{"meta":{"query_hash":"180aa56179b7","filters":{"topic":"Software Engineering Research"},"cohort_total":3468,"direct_labels_cover":10,"predictions_cover":3468,"exported":3468,"export_cap":100000,"truncated":false,"label_status":"direct model label, unvalidated","prediction_status":"machine_predicted_unvalidated (Codex and Gemma teacher distillation)","score_status":"score_only:v0-immature-baseline","snapshot":{"source":"OpenAlex, pinned release, all 482 partitions","release":"2026-06-24","frame_built":"2026-07-12"},"permalink":"https://metacan.xera.ac/q/180aa56179b7","api":"https://metacan.xera.ac/api/v1/cohort?topic=Software+Engineering+Research"},"results":[{"id":"W100770364","doi":"","title":"Linker-Based Program Extraction and Its Uses in Studying Software Evolution","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Software; Extractor; Source code; Software system; Programming language; Code (set theory); Backporting; Software development; Software construction; Theoretical computer science; Software engineering; Engineering","score_opus":0.03662086625094669,"score_gpt":0.31527876419812445,"score_spread":0.27865789794717777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W100770364","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2171296,0.001309897,0.77301484,0.00038929994,0.000013592426,0.0001689231,0.000836812,0.0028757527,0.0042612073],"genre_scores_gemma":[0.4182046,0.00084656815,0.57872087,0.000053429183,0.000018932782,0.00016492006,0.0010198127,0.0002221098,0.0007486614],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99861264,0.0006967015,0.00009132984,0.00024520184,0.00031902065,0.000035031066],"domain_scores_gemma":[0.9752641,0.019595353,0.0018360247,0.0021980603,0.0009154049,0.00019111464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002803625,0.0004284042,0.0004008468,0.007118372,0.0005648023,0.00087100704,0.00061306276,0.0005674295,0.0010490753],"category_scores_gemma":[0.0187436,0.00040485722,0.0004890964,0.0067696422,0.0012205486,0.0027220477,0.00093405356,0.00077831437,0.00021119758],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002420774,0.0003577085,0.13041237,0.0009738378,0.00020730811,0.00063281297,0.00528259,0.04431527,0.033208016,0.0472818,0.0015109015,0.7355753],"study_design_scores_gemma":[0.0001014277,0.0008680757,0.17047228,0.00030969712,0.00037410765,0.0033864789,0.0017821313,0.55971205,0.09440395,0.1207123,0.047644164,0.0002333675],"about_ca_topic_score_codex":0.0023266398,"about_ca_topic_score_gemma":0.00277466,"teacher_disagreement_score":0.007118372,"about_ca_system_score_codex":0.000624829,"about_ca_system_score_gemma":0.0006512971,"threshold_uncertainty_score":0.014827132},"labels":[],"label_agreement":null},{"id":"W100954132","doi":"10.1007/978-3-642-54092-9_13","title":"Investigating the Applicability of the Laws of Software Evolution: A Metrics Based Study","year":2013,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Software evolution; Metric (unit); Computer science; Software; Java; Software metric; Quality (philosophy); Empirical research; Open source software; Object (grammar); Software quality; Software development; Software engineering; Data science; Software construction; Programming language; Artificial intelligence; Mathematics; Engineering; Operations management; Statistics","score_opus":0.051741168538568034,"score_gpt":0.29699221642596985,"score_spread":0.24525104788740182,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W100954132","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53389335,0.010403622,0.3987186,0.005505932,0.00015334791,0.0002594652,0.00015873105,0.00021509209,0.050691813],"genre_scores_gemma":[0.9482891,0.0030559078,0.04665571,0.00017015777,0.000074725074,0.00013631965,0.00009909136,0.000079759026,0.001439309],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9902199,0.0064183064,0.00044296368,0.00048140765,0.002242474,0.00019488181],"domain_scores_gemma":[0.81556875,0.16534737,0.007095548,0.0053808466,0.005887068,0.00072043],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010937081,0.0008003665,0.00066392537,0.0039265836,0.0006409358,0.002482377,0.0018591098,0.0015107119,0.0018059012],"category_scores_gemma":[0.14301825,0.00043344242,0.00088542135,0.0057679336,0.0032333895,0.009311139,0.0017371965,0.00213338,0.00020211999],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000116817035,0.00034418874,0.051752325,0.0006629266,0.0002297691,0.0003038202,0.0031610823,0.08110923,0.0023879996,0.6191541,0.002519874,0.23825796],"study_design_scores_gemma":[0.00003388862,0.00058591884,0.035286643,0.00046569647,0.0001235281,0.00044584088,0.0019918627,0.3765801,0.001834172,0.57552457,0.007073056,0.000054727352],"about_ca_topic_score_codex":0.0027681645,"about_ca_topic_score_gemma":0.0016643878,"teacher_disagreement_score":0.9890629,"about_ca_system_score_codex":0.0019687833,"about_ca_system_score_gemma":0.0013154492,"threshold_uncertainty_score":0.05784148},"labels":[],"label_agreement":null},{"id":"W10134666","doi":"","title":"Managing Software Reuse with Perforce","year":2003,"lang":"en","type":"article","venue":"Canadian journal of medical technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Reuse; Software; Software engineering; Engineering; Programming language","score_opus":0.010069548407787003,"score_gpt":0.23589781236582352,"score_spread":0.2258282639580365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W10134666","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3328743,0.0011494926,0.6251672,0.002260851,0.000093431474,0.0002763423,0.00007038776,0.007164159,0.030943789],"genre_scores_gemma":[0.7660488,0.00036994636,0.22452395,0.00021806612,0.000060002283,0.00011423784,0.0001844573,0.0007628958,0.0077176313],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9941444,0.001448685,0.00042211317,0.0007152249,0.0025175824,0.00075207546],"domain_scores_gemma":[0.976014,0.0077082505,0.0018386513,0.010285963,0.0034119147,0.0007411323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005713966,0.0007712793,0.0007772978,0.003411166,0.0014817324,0.0040471912,0.0025455507,0.0015074542,0.0033307406],"category_scores_gemma":[0.035717223,0.0007682966,0.0009863611,0.0027548976,0.0018690039,0.010990787,0.0062964056,0.0020109075,0.000926143],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002825163,0.0005847586,0.018850388,0.00024960382,0.00020814389,0.00060260086,0.0020922255,0.03651484,0.0130464155,0.052515417,0.0054432377,0.8696098],"study_design_scores_gemma":[0.00018086484,0.00091322494,0.009919644,0.00018576128,0.0005793357,0.0021025566,0.0024337783,0.6067454,0.048082676,0.28022757,0.048480883,0.00014823477],"about_ca_topic_score_codex":0.0030718835,"about_ca_topic_score_gemma":0.0043888576,"teacher_disagreement_score":0.005713966,"about_ca_system_score_codex":0.0013286776,"about_ca_system_score_gemma":0.0030923279,"threshold_uncertainty_score":0.03021872},"labels":[],"label_agreement":null},{"id":"W101628381","doi":"","title":"Exploring an Open Source Data Mining Environment for Software Product Quality Decision Making","year":2006,"lang":"en","type":"article","venue":"Joint Conference on Knowledge-Based Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Software quality; Quality (philosophy); Software quality control; Software metric; Software; Software measurement; Software engineering; Data mining; Software development","score_opus":0.22859666811365403,"score_gpt":0.3480682565296247,"score_spread":0.11947158841597066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W101628381","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07117177,0.0002311421,0.90411836,0.0008831006,0.000052302898,0.0005084921,0.0028531288,0.018678881,0.0015028252],"genre_scores_gemma":[0.18249573,0.00013035163,0.8111198,0.000102721184,0.00003471245,0.00065759243,0.004423903,0.00045876912,0.000576398],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949942,0.0020562983,0.00067782233,0.0007556209,0.0013848783,0.0001311117],"domain_scores_gemma":[0.9498272,0.03867799,0.0017868639,0.0053123287,0.003825141,0.00057039125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010636386,0.0009960571,0.0008328584,0.0035422812,0.0010305886,0.0028673345,0.0021371513,0.0011501833,0.0015160902],"category_scores_gemma":[0.052864473,0.0006857168,0.0011619992,0.004818104,0.00060199224,0.0052010594,0.002995676,0.0021930889,0.0006682863],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001919874,0.00304736,0.047198385,0.0010130253,0.00058785494,0.0010583131,0.0030009823,0.17757872,0.010129601,0.031473428,0.0125544695,0.71043795],"study_design_scores_gemma":[0.0001375856,0.00018028682,0.0043666256,0.00008461945,0.00006858851,0.00019827769,0.00025502904,0.9512381,0.008832981,0.023814585,0.010774354,0.0000490885],"about_ca_topic_score_codex":0.0040696408,"about_ca_topic_score_gemma":0.005184787,"teacher_disagreement_score":0.010636386,"about_ca_system_score_codex":0.0009343614,"about_ca_system_score_gemma":0.0017565048,"threshold_uncertainty_score":0.056251287},"labels":[],"label_agreement":null},{"id":"W1024906986","doi":"10.1007/s10664-015-9387-3","title":"A contextual approach towards more accurate duplicate bug report detection and ranking","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":90,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Data deduplication; Software bug; Android (operating system); Software; Ranking (information retrieval); Context (archaeology); Data science; Software engineering; Data mining; World Wide Web; Information retrieval; Database","score_opus":0.05477320630831465,"score_gpt":0.3061968965378645,"score_spread":0.25142369022954986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1024906986","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046207752,0.0019917546,0.93771666,0.0008581258,0.00032150486,0.00038491536,0.0018881423,0.0071989605,0.0034321365],"genre_scores_gemma":[0.29234838,0.0005453317,0.70032775,0.00045024185,0.00053776824,0.00030801856,0.002719136,0.00053159177,0.0022317164],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98705286,0.003792562,0.0012008535,0.0029276395,0.004370277,0.0006557617],"domain_scores_gemma":[0.96314037,0.011398085,0.0038102854,0.009844602,0.010938525,0.00086806295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055445614,0.0016466171,0.0030921355,0.009286934,0.0018235494,0.0041617835,0.0028754016,0.002600113,0.0044913343],"category_scores_gemma":[0.03724048,0.00093015924,0.0014202576,0.007033832,0.00093471457,0.0041001844,0.004323479,0.0022933108,0.002688098],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011199029,0.0012526659,0.03931684,0.0013809397,0.00050587236,0.00068157335,0.0010973972,0.02977865,0.067421034,0.019268462,0.025200605,0.81297606],"study_design_scores_gemma":[0.00026428766,0.0013374211,0.032803357,0.00043058503,0.00093393447,0.0018495712,0.0012586724,0.7933056,0.06508187,0.05318017,0.049160376,0.00039421127],"about_ca_topic_score_codex":0.005432454,"about_ca_topic_score_gemma":0.016894316,"teacher_disagreement_score":0.009286934,"about_ca_system_score_codex":0.00084165606,"about_ca_system_score_gemma":0.0041713812,"threshold_uncertainty_score":0.029322803},"labels":[],"label_agreement":null},{"id":"W102554878","doi":"","title":"Modeling Human Aspects to Enhance Software Quality Management","year":2012,"lang":"en","type":"article","venue":"ENLIGHTEN (Jurnal Bimbingan dan Konseling Islam)","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software quality; Quality (philosophy); Software; Software engineering; Software development; Operating system","score_opus":0.030422072997987694,"score_gpt":0.32748871232229576,"score_spread":0.29706663932430805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W102554878","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13631243,0.0017395604,0.82624406,0.0029512038,0.00012106178,0.00026181736,0.00023059512,0.00034872495,0.031790566],"genre_scores_gemma":[0.8690237,0.0012129672,0.12278645,0.00014883332,0.00006900079,0.00020601206,0.00017084506,0.000057080942,0.0063251033],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993507,0.00034346012,0.000027840586,0.00009610504,0.0001261625,0.000055677323],"domain_scores_gemma":[0.99737525,0.0017402386,0.0003357014,0.0002017662,0.00025835962,0.0000887104],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012980168,0.0005294787,0.00032627556,0.0007612787,0.00036404215,0.0016655744,0.0007228708,0.0008326847,0.0035724589],"category_scores_gemma":[0.0060298946,0.0002536836,0.0005763851,0.0007364489,0.00061506225,0.0018209092,0.00090616936,0.00064470235,0.00038379894],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000075528005,0.00033493558,0.018546637,0.00026891546,0.00013504417,0.00018777947,0.0010331131,0.749111,0.0031075007,0.131761,0.0019930857,0.0934455],"study_design_scores_gemma":[0.000017829809,0.00009274845,0.0028629573,0.000043599528,0.000050819726,0.000035888286,0.00025266423,0.91933006,0.0005750932,0.07125561,0.0054664994,0.000016234286],"about_ca_topic_score_codex":0.007916323,"about_ca_topic_score_gemma":0.006100808,"teacher_disagreement_score":0.007916323,"about_ca_system_score_codex":0.0011061392,"about_ca_system_score_gemma":0.0011575265,"threshold_uncertainty_score":0.015740514},"labels":[],"label_agreement":null},{"id":"W104392493","doi":"","title":"Proceedings of the 17th ACM SIGPLAN conference on Object-oriented programming, systems, languages, and applications","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Presentation (obstetrics); Session (web analytics); Object (grammar); Library science; Database; World Wide Web; Artificial intelligence","score_opus":0.02377048046847855,"score_gpt":0.2690139825263813,"score_spread":0.24524350205790277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W104392493","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014163599,0.09685979,0.298273,0.02421748,0.086177066,0.0032154054,0.008044215,0.019653434,0.44939598],"genre_scores_gemma":[0.006844151,0.041118156,0.0680518,0.0042951163,0.00446768,0.0009285866,0.011890249,0.0042917347,0.8581126],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983327,0.0004156729,0.00014239918,0.00033008296,0.0006538691,0.00012528774],"domain_scores_gemma":[0.99637765,0.0009792906,0.00020255911,0.00049653853,0.0010248846,0.0009190898],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00351624,0.0026160097,0.0022195925,0.0014055182,0.0014154162,0.007020815,0.002030814,0.0019738672,0.19573472],"category_scores_gemma":[0.005373417,0.0015630763,0.0013776316,0.0018112409,0.0010553132,0.0059170597,0.0027212868,0.0064950846,0.13697179],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016091397,0.000257573,0.00042372587,0.00038654415,0.000046821937,0.00016850264,0.0002474057,0.00039031176,0.0021394838,0.003330238,0.78787625,0.20457236],"study_design_scores_gemma":[0.000037078007,0.000060831855,0.0005625131,0.00021999115,0.000026605365,0.00017999456,0.00013360761,0.0007297522,0.00031033639,0.0018707765,0.9958502,0.000018388304],"about_ca_topic_score_codex":0.005259959,"about_ca_topic_score_gemma":0.009058769,"teacher_disagreement_score":0.19573472,"about_ca_system_score_codex":0.0010792685,"about_ca_system_score_gemma":0.0034680858,"threshold_uncertainty_score":0.6547979},"labels":[],"label_agreement":null},{"id":"W10528201","doi":"","title":"A Comparative Study of Attribute Weighting Techniques for Software Defect Prediction Using Case-based Reasoning.","year":2010,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Weighting; Case-based reasoning; Artificial intelligence; Software; Data mining; Software bug; Machine learning; Programming language","score_opus":0.023671926535258184,"score_gpt":0.28584333824540004,"score_spread":0.26217141171014186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W10528201","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053188544,0.0020308918,0.93873435,0.0004647929,0.00008901939,0.0004953823,0.0004971409,0.002045831,0.0024540762],"genre_scores_gemma":[0.3202675,0.0008011702,0.676991,0.00011200463,0.000054669737,0.0002190386,0.0008260125,0.000093616756,0.00063506945],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9938261,0.0025936947,0.0007740622,0.00085264584,0.0017377312,0.00021570294],"domain_scores_gemma":[0.9668229,0.028011251,0.0010496257,0.001527954,0.002095576,0.00049274554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010746497,0.0014427203,0.0013157749,0.007901805,0.0005730742,0.0019922717,0.0021668475,0.0015734701,0.00245918],"category_scores_gemma":[0.028041525,0.00056618295,0.0020580394,0.0038604853,0.00063129875,0.0031054444,0.0014588264,0.001443441,0.0005027767],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064441166,0.00092111994,0.01486484,0.0005764386,0.00061534916,0.00025435715,0.00035881158,0.16717541,0.0025849727,0.009678983,0.0026446057,0.7996807],"study_design_scores_gemma":[0.00003417277,0.0000868887,0.0011649702,0.000048485297,0.00008289421,0.00010307253,0.00008962052,0.9912703,0.0006477419,0.0057160794,0.0007391828,0.00001662019],"about_ca_topic_score_codex":0.010392766,"about_ca_topic_score_gemma":0.00735849,"teacher_disagreement_score":0.010746497,"about_ca_system_score_codex":0.0013455375,"about_ca_system_score_gemma":0.0017500302,"threshold_uncertainty_score":0.056833565},"labels":[],"label_agreement":null},{"id":"W107899461","doi":"10.71781/11002","title":"Migrating legacy system towards object technology","year":2005,"lang":"en","type":"dissertation","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal","funders":"","keywords":"Legacy system; Business process reengineering; Computer science; Legacy code; Software engineering; Software system; Software maintenance; Software evolution; Software; Systems engineering; Engineering; Software construction; Programming language; Operations management","score_opus":0.023848519487259714,"score_gpt":0.32442944025794274,"score_spread":0.30058092077068305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W107899461","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06659482,0.09101231,0.2859493,0.07856076,0.017739618,0.0008575428,0.0067963395,0.01930421,0.4331851],"genre_scores_gemma":[0.10373052,0.063829854,0.079094976,0.0032836453,0.0026072504,0.00020174331,0.008419406,0.0027751282,0.7360575],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998226,0.0002221006,0.0001123642,0.00018198791,0.0009883245,0.00026932568],"domain_scores_gemma":[0.9962102,0.00045636948,0.000098186014,0.00070916524,0.0021724526,0.00035362897],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029957586,0.0005950048,0.0005610944,0.0029338258,0.0018783397,0.0069398675,0.0014293656,0.0011914734,0.034014557],"category_scores_gemma":[0.005231117,0.00055814785,0.0006248995,0.003744147,0.0008578153,0.0044161617,0.0017212394,0.0013425683,0.007629459],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011111388,0.00016576394,0.0030212083,0.0005895043,0.000041799125,0.0007363455,0.0017531813,0.0018306675,0.01502322,0.06871664,0.2981047,0.60990584],"study_design_scores_gemma":[0.000023479204,0.000043045307,0.004914277,0.00025019798,0.000025471221,0.00026756324,0.00036059014,0.0025048638,0.0056552375,0.0061786408,0.97974193,0.000034725963],"about_ca_topic_score_codex":0.22446299,"about_ca_topic_score_gemma":0.23379378,"teacher_disagreement_score":0.22446299,"about_ca_system_score_codex":0.0057119187,"about_ca_system_score_gemma":0.008170723,"threshold_uncertainty_score":0.4463129},"labels":[],"label_agreement":null},{"id":"W10897080","doi":"10.1002/1098-108x(200009)28:2<181::aid-eat7>3.0.co;2-k","title":"Detecting model refactoring opportunities using heuristic search","year":2011,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; École de Technologie Supérieure","funders":"","keywords":"Code refactoring; Computer science; Model-driven architecture; Heuristic; Process (computing); Software engineering; Artificial intelligence; Software development; Software; Programming language","score_opus":0.5083945867567765,"score_gpt":0.4311035929333741,"score_spread":0.07729099382340243,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W10897080","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5081077,0.003460436,0.4608897,0.0015971431,0.00016465764,0.0012330018,0.003874614,0.014322181,0.006350583],"genre_scores_gemma":[0.60706234,0.0005780639,0.38464752,0.00031779907,0.00005037986,0.00027971558,0.005567973,0.0003593853,0.0011368342],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965815,0.0010728716,0.00033623684,0.00059118087,0.0011431305,0.00027515338],"domain_scores_gemma":[0.9713794,0.02298583,0.0019971067,0.0014686235,0.0017039386,0.00046500066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034164074,0.0022791808,0.0014459726,0.007173145,0.000892326,0.0022332773,0.0028598337,0.0019695433,0.0024377124],"category_scores_gemma":[0.029834395,0.0007880282,0.002072036,0.0030763918,0.00064140634,0.0023369652,0.0015469358,0.0012815398,0.00059721014],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026566265,0.0021210045,0.14855376,0.0025011187,0.0013116773,0.0048130048,0.001938696,0.15727232,0.023545526,0.011258122,0.019734103,0.624294],"study_design_scores_gemma":[0.0004201599,0.0006129826,0.009504801,0.0002686437,0.00063099345,0.0013345195,0.001169576,0.95449567,0.010838247,0.015711881,0.004894729,0.000117818374],"about_ca_topic_score_codex":0.011848731,"about_ca_topic_score_gemma":0.027499821,"teacher_disagreement_score":0.011848731,"about_ca_system_score_codex":0.0011181505,"about_ca_system_score_gemma":0.0034630795,"threshold_uncertainty_score":0.02355951},"labels":[],"label_agreement":null},{"id":"W112868136","doi":"","title":"Proceedings of the 2nd Workshop on Managing Technical Debt","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Debt; Engineering management; Brainstorming; Maintainability; Computer science; Software; Engineering; Software engineering; Engineering ethics; Software development; Business; Finance","score_opus":0.03162559602200429,"score_gpt":0.25447223511431016,"score_spread":0.22284663909230587,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W112868136","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013419405,0.04747303,0.11108348,0.13949387,0.22704133,0.0012840994,0.003354673,0.0033314852,0.45351863],"genre_scores_gemma":[0.036183972,0.016799983,0.032673083,0.009170957,0.018069733,0.0007713759,0.0051236805,0.0023755901,0.87883157],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99649954,0.0009370512,0.00022841182,0.00063317316,0.0011563359,0.0005454519],"domain_scores_gemma":[0.99336714,0.0013001863,0.00023538,0.000665277,0.0023071563,0.0021247894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058854274,0.0013300158,0.0009982436,0.0016176384,0.0027070802,0.009245946,0.0024620665,0.0038203602,0.12241189],"category_scores_gemma":[0.010298663,0.0007277421,0.0015751461,0.0014519175,0.0011737641,0.007726682,0.006165105,0.005814734,0.045099262],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001534135,0.00012967417,0.0002732954,0.00031321656,0.000018841327,0.00019267335,0.0010793066,0.00048426716,0.0018148453,0.008954791,0.8498907,0.13669509],"study_design_scores_gemma":[0.000012194362,0.000038308866,0.000326487,0.00017702406,0.000008950766,0.00009563841,0.00043932663,0.00023195727,0.00041810467,0.0033208146,0.9949149,0.00001627931],"about_ca_topic_score_codex":0.00243256,"about_ca_topic_score_gemma":0.0059751114,"teacher_disagreement_score":0.12241189,"about_ca_system_score_codex":0.0022692776,"about_ca_system_score_gemma":0.004005677,"threshold_uncertainty_score":0.4095086},"labels":[],"label_agreement":null},{"id":"W115256144","doi":"10.11575/prism/35560","title":"Connectivity of co-changed method groups: a case study on open source systems","year":2012,"lang":"en","type":"article","venue":"PRISM (University of Calgary)","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Java; Source code; Source lines of code; Cloning (programming); Open source software; Software development; Software maintenance; Software engineering; Software; Software system; Open source; Coding (social sciences); Code review; Codebase; Programming language; Static program analysis","score_opus":0.045522582954166474,"score_gpt":0.29647186121588426,"score_spread":0.25094927826171776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W115256144","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9931572,0.00013668653,0.005194598,0.00026890144,0.000006824055,0.00009346857,0.000120916484,0.00009115199,0.00093020656],"genre_scores_gemma":[0.98369366,0.00014941754,0.01491792,0.00007101913,0.000014264818,0.00007645,0.0002925275,0.00006980289,0.0007148872],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9940281,0.0028098258,0.0003511782,0.0007408708,0.0016564352,0.00041350423],"domain_scores_gemma":[0.9055471,0.074259184,0.00664171,0.0056317756,0.005403103,0.0025171707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050223484,0.0005186677,0.00048003608,0.0027123063,0.0025700282,0.0015164164,0.002086856,0.0022262416,0.0009044045],"category_scores_gemma":[0.029727727,0.00045483548,0.0008282693,0.0029208101,0.0022626712,0.0025865347,0.0019785366,0.0016601448,0.00018017225],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011507882,0.0050470377,0.53304553,0.0018691466,0.00043846265,0.057130728,0.082677744,0.07128285,0.035492536,0.012318552,0.006277217,0.19326949],"study_design_scores_gemma":[0.0006287358,0.0041541387,0.53461874,0.00056258775,0.00055759784,0.024721377,0.06707784,0.2637506,0.045120865,0.017650234,0.04075697,0.0004003725],"about_ca_topic_score_codex":0.012796925,"about_ca_topic_score_gemma":0.022223352,"teacher_disagreement_score":0.012796925,"about_ca_system_score_codex":0.0016872085,"about_ca_system_score_gemma":0.0011704598,"threshold_uncertainty_score":0.026561022},"labels":[],"label_agreement":null},{"id":"W116782570","doi":"10.1007/978-0-387-21599-0_7","title":"Approaches to Clustering for Program Comprehension and Remodularization","year":2002,"lang":"en","type":"book-chapter","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Program comprehension; Cluster analysis; Computer science; Comprehension; Artificial intelligence; Programming language; Software","score_opus":0.05777975468218309,"score_gpt":0.2594115183847261,"score_spread":0.201631763702543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W116782570","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000923223,0.00030656313,0.993522,0.00020635368,0.000024142671,0.000040771483,0.00006596655,0.00088954606,0.0040214523],"genre_scores_gemma":[0.032640353,0.0006458131,0.9556556,0.00013222599,0.00011310267,0.00019674373,0.00056077924,0.000855858,0.009199634],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99732447,0.00085230486,0.0001893961,0.00062690885,0.0008346282,0.00017224497],"domain_scores_gemma":[0.99382174,0.0030163906,0.00023607278,0.0018322904,0.0009836225,0.000109742374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026509622,0.0015488687,0.0017285831,0.003739154,0.0019173129,0.0030933411,0.005246722,0.002013492,0.011920446],"category_scores_gemma":[0.011820036,0.0011151986,0.0029057402,0.0056967624,0.0029112024,0.007868767,0.0032158422,0.00326153,0.003323971],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003689692,0.000079214056,0.00040201258,0.00030879362,0.00007146743,0.00005899317,0.00086971256,0.020373153,0.00161135,0.597298,0.01646334,0.362427],"study_design_scores_gemma":[0.000011074735,0.000018443594,0.00028429975,0.000074372714,0.000043925833,0.000127693,0.00027714422,0.13545392,0.00251732,0.82707286,0.034087166,0.000031860232],"about_ca_topic_score_codex":0.004163238,"about_ca_topic_score_gemma":0.0069577345,"teacher_disagreement_score":0.011920446,"about_ca_system_score_codex":0.002471513,"about_ca_system_score_gemma":0.0013928495,"threshold_uncertainty_score":0.03987789},"labels":[],"label_agreement":null},{"id":"W1171898901","doi":"10.1007/978-3-319-17837-0_4","title":"Implicit Coordination: A Case Study of the Rails OSS Project","year":2015,"lang":"en","type":"book-chapter","venue":"IFIP advances in information and communication technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Coordination game; Computer science; Knowledge management; Coordination complex; Mathematics; Chemistry","score_opus":0.025973261056593186,"score_gpt":0.3193726787126176,"score_spread":0.2933994176560244,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1171898901","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92304003,0.0007490536,0.015043173,0.006395811,0.00009986552,0.00027679486,0.0002937877,0.00014063795,0.053960804],"genre_scores_gemma":[0.943717,0.0012541966,0.022118883,0.0009482697,0.00006586798,0.00024824674,0.00028812233,0.00019090348,0.031168532],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9938034,0.004022639,0.0001989555,0.00041118875,0.0010368099,0.00052702276],"domain_scores_gemma":[0.99186,0.0050624763,0.0006643482,0.0006088363,0.0005968292,0.0012074208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004618731,0.0005537057,0.00044261175,0.0014304646,0.00896507,0.003287217,0.0018175651,0.0043067327,0.004072696],"category_scores_gemma":[0.009985894,0.00037458222,0.00043743884,0.0021992214,0.00440122,0.0029006288,0.0042306413,0.0027108868,0.0009907646],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003788836,0.0019534857,0.03409899,0.00077677687,0.00006866596,0.09600137,0.65157443,0.0032539703,0.006733451,0.0630852,0.029831903,0.11224286],"study_design_scores_gemma":[0.00014446296,0.0011777378,0.034184128,0.0008080771,0.00005653583,0.029425474,0.5926443,0.0062134173,0.003983037,0.013532298,0.31766865,0.00016202925],"about_ca_topic_score_codex":0.012298342,"about_ca_topic_score_gemma":0.034271546,"teacher_disagreement_score":0.012298342,"about_ca_system_score_codex":0.002657262,"about_ca_system_score_gemma":0.0023086679,"threshold_uncertainty_score":0.02445352},"labels":[],"label_agreement":null},{"id":"W1173311655","doi":"10.1016/j.infsof.2015.07.007","title":"Learning dependency-based change impact predictors using independent change histories","year":2015,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Change impact analysis; Variety (cybernetics); Software; Machine learning; Precision and recall; Context (archaeology); Artificial intelligence; Data mining; Dependency (UML); Software system; Data science","score_opus":0.04894197911269241,"score_gpt":0.2868627333318484,"score_spread":0.237920754219156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1173311655","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70933306,0.0015265926,0.27552494,0.0007257503,0.00022727116,0.00023672792,0.0049373615,0.003373503,0.0041147857],"genre_scores_gemma":[0.9678772,0.0002901191,0.02558602,0.000051583786,0.000098703254,0.00012280843,0.0044812923,0.00009837913,0.001393935],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998838,0.0002393817,0.000104660736,0.00044125546,0.00024986704,0.00012690887],"domain_scores_gemma":[0.97780734,0.016787352,0.001439199,0.0013671045,0.0019969058,0.00060216454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027293349,0.00103362,0.0010139052,0.0054114303,0.0005233004,0.0011540046,0.001165096,0.0010623119,0.0029877792],"category_scores_gemma":[0.017496781,0.00059503806,0.0010165892,0.0029742366,0.00038727763,0.0028202315,0.0010661519,0.002334996,0.001662064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011631126,0.0015656115,0.47997075,0.00025993405,0.000473323,0.00028556993,0.00030009155,0.08717368,0.0029336256,0.0022116215,0.009274382,0.41438833],"study_design_scores_gemma":[0.00005126101,0.0003336744,0.050007306,0.000057746267,0.00017842368,0.00010291234,0.00012890114,0.93981457,0.0019872766,0.005944927,0.0013513233,0.000041719282],"about_ca_topic_score_codex":0.006742345,"about_ca_topic_score_gemma":0.013242534,"teacher_disagreement_score":0.006742345,"about_ca_system_score_codex":0.00061356515,"about_ca_system_score_gemma":0.0012350779,"threshold_uncertainty_score":0.014434338},"labels":[],"label_agreement":null},{"id":"W127903220","doi":"","title":"A Design for Evidence-based Software Architecture Research","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Software engineering; Data science; Architecture; Software architecture; Process (computing); Empirical evidence; Software; Management science; Engineering; Programming language","score_opus":0.1559911719529537,"score_gpt":0.3692105742918407,"score_spread":0.21321940233888698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W127903220","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008404655,0.009762642,0.57716066,0.03653422,0.0054616095,0.312628,0.003331705,0.0008217607,0.045894805],"genre_scores_gemma":[0.013786504,0.001763707,0.73213834,0.003483299,0.00018546614,0.24654415,0.0003621719,0.00007377481,0.0016625674],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.6021282,0.31122872,0.044985,0.01615675,0.022082908,0.0034184284],"domain_scores_gemma":[0.5754834,0.30489045,0.017300488,0.045114752,0.048016492,0.009194382],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.26762033,0.0021290267,0.0062630903,0.025157068,0.007715248,0.018467503,0.0076636183,0.014042558,0.026904738],"category_scores_gemma":[0.31569746,0.0038686136,0.006742035,0.019094104,0.011642575,0.017530387,0.018677505,0.011456775,0.0060986113],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026071805,0.0010824823,0.0036061946,0.044043254,0.0009704017,0.0013591237,0.014846577,0.0022694974,0.0025039306,0.6251463,0.02483521,0.27672982],"study_design_scores_gemma":[0.0072720875,0.0039174263,0.0038956506,0.05052649,0.002428481,0.0014142669,0.014490195,0.0042520766,0.0031211737,0.3339407,0.5743234,0.00041808997],"about_ca_topic_score_codex":0.0011017675,"about_ca_topic_score_gemma":0.001848813,"teacher_disagreement_score":0.7323797,"about_ca_system_score_codex":0.01573192,"about_ca_system_score_gemma":0.03225083,"threshold_uncertainty_score":0.90315455},"labels":[],"label_agreement":null},{"id":"W128336448","doi":"10.5555/2555523.2555553","title":"Revisiting prior empirical findings for mobile apps: an empirical case study on the 15 most popular open-source Android apps","year":2013,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Android (operating system); Computer science; Software; Mobile device; Mobile apps; World Wide Web; Mobile computing; Empirical research; Operating system; Multimedia","score_opus":0.03552042791183693,"score_gpt":0.3201485272367965,"score_spread":0.2846280993249596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W128336448","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.988853,0.00096667826,0.0011577847,0.0012545796,0.000021186423,0.00021180777,0.0002622175,0.000012356941,0.007260389],"genre_scores_gemma":[0.99453956,0.001122414,0.002077929,0.0003643665,0.000021554459,0.0001967356,0.0003652131,0.000039361716,0.0012729011],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9752256,0.009116911,0.0025294204,0.0025848073,0.009282967,0.001260375],"domain_scores_gemma":[0.6448645,0.28433183,0.026005942,0.011210407,0.03114292,0.0024443553],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.025096739,0.00064018945,0.0005951945,0.005782582,0.0030961507,0.0056345863,0.002477285,0.001862626,0.003684494],"category_scores_gemma":[0.15491949,0.00070939725,0.00062447175,0.0066375155,0.0050271302,0.0083275605,0.003224488,0.0032388929,0.0009861623],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021560596,0.0028254983,0.6506489,0.0019815282,0.00010443886,0.005111664,0.25149763,0.0005931944,0.001706633,0.007656915,0.0057137636,0.07194428],"study_design_scores_gemma":[0.00004529599,0.0006099972,0.588884,0.0026333604,0.0001472108,0.0025578719,0.36558333,0.003399797,0.002074113,0.0018681061,0.03209929,0.000097641256],"about_ca_topic_score_codex":0.01496424,"about_ca_topic_score_gemma":0.02784616,"teacher_disagreement_score":0.9749033,"about_ca_system_score_codex":0.0030868733,"about_ca_system_score_gemma":0.0032948267,"threshold_uncertainty_score":0.13272583},"labels":[],"label_agreement":null},{"id":"W128739217","doi":"10.1007/978-1-4615-0457-3_5","title":"Designing a Software Exploration Tool Using a Cognitive Framework","year":2003,"lang":"en","type":"book-chapter","venue":"Software Visualization","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software visualization; Computer science; Visualization; Debugging; Software engineering; Software; Documentation; Program comprehension; Profiling (computer programming); Software construction; Software framework; Software development; Human–computer interaction; World Wide Web; Software system; Programming language; Artificial intelligence","score_opus":0.05997051868444286,"score_gpt":0.3217217251836076,"score_spread":0.2617512064991647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W128739217","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031220932,0.00031500685,0.98101634,0.000143555,0.00003514238,0.00004258778,0.000021455264,0.0013483108,0.013955562],"genre_scores_gemma":[0.03014681,0.000494854,0.95314294,0.000058424583,0.000012927907,0.000108831766,0.0000838425,0.000595189,0.015356217],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996531,0.00008741191,0.000018330187,0.0000730436,0.00012774445,0.00004039498],"domain_scores_gemma":[0.9992618,0.0005074151,0.000017348177,0.000076435485,0.000090898466,0.000046065892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00071821763,0.0013106313,0.0005488445,0.0013742142,0.0008649403,0.003326604,0.0015052259,0.0011688481,0.008140856],"category_scores_gemma":[0.0020796827,0.00084141234,0.0011589432,0.0011668651,0.0012827036,0.004705853,0.0014892096,0.0016545248,0.0019932284],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000051347684,0.000100924466,0.0005150735,0.0007433735,0.000048996626,0.0005124786,0.0071185194,0.0098152915,0.04007142,0.2681319,0.015801651,0.65708905],"study_design_scores_gemma":[0.00007129198,0.00026035774,0.0010526896,0.0008846581,0.0002097373,0.003135718,0.0030472523,0.16194971,0.044325076,0.336199,0.4486803,0.00018417514],"about_ca_topic_score_codex":0.0014353954,"about_ca_topic_score_gemma":0.0024284406,"teacher_disagreement_score":0.008140856,"about_ca_system_score_codex":0.0005225751,"about_ca_system_score_gemma":0.0009350486,"threshold_uncertainty_score":0.027233899},"labels":[],"label_agreement":null},{"id":"W131477820","doi":"","title":"A source-to-source transformation tool for error fixing","year":2013,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus","funders":"","keywords":"Source code; Computer science; Open source; Program transformation; Transformation (genetics); Crash; Code (set theory); Operating system; Programming language; Computer engineering; Software; Set (abstract data type)","score_opus":0.08341426590995744,"score_gpt":0.3805953589430776,"score_spread":0.2971810930331202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W131477820","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017979449,0.00007447425,0.9219351,0.00008887443,0.0000761318,0.00021284865,0.00028866917,0.07485176,0.0006741197],"genre_scores_gemma":[0.058151048,0.0002203356,0.92004824,0.00017866922,0.00006571024,0.0005156479,0.0019249113,0.015523384,0.0033720415],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9939307,0.0010105812,0.0006384995,0.0009359129,0.0032046216,0.000279767],"domain_scores_gemma":[0.9830913,0.006376056,0.0014604307,0.004388491,0.0043646833,0.00031909585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00369623,0.0020017752,0.0010565727,0.0035678146,0.00091812667,0.0017617381,0.003148696,0.0020165585,0.009269353],"category_scores_gemma":[0.022205908,0.0012241598,0.0015325746,0.0017998496,0.0011134057,0.0023247502,0.00285197,0.0033356068,0.006577839],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041667087,0.0005881983,0.0045187697,0.0014334376,0.00026134594,0.002078853,0.0013917442,0.022172369,0.089973144,0.016678248,0.04917797,0.8113094],"study_design_scores_gemma":[0.00033422103,0.00062753464,0.0046148393,0.0006044618,0.00030410488,0.00477956,0.00033821136,0.3875881,0.34150195,0.025641179,0.23322998,0.0004359393],"about_ca_topic_score_codex":0.0011737993,"about_ca_topic_score_gemma":0.00088348176,"teacher_disagreement_score":0.009269353,"about_ca_system_score_codex":0.0005160253,"about_ca_system_score_gemma":0.0024829805,"threshold_uncertainty_score":0.031009078},"labels":[],"label_agreement":null},{"id":"W1419069490","doi":"10.4018/978-1-4666-4785-5.ch014","title":"Improving the Performance of Neuro-Fuzzy Function Point Backfiring Model with Additional Environmental Factors","year":2013,"lang":"en","type":"book-chapter","venue":"Advances in computational intelligence and robotics book series","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Fuzzy logic; Point (geometry); Computer science; Artificial neural network; Function (biology); Function point; Code (set theory); Software; Neuro-fuzzy; Mathematical optimization; Algorithm; Artificial intelligence; Mathematics; Software development; Fuzzy control system","score_opus":0.014925922676088004,"score_gpt":0.215956003084556,"score_spread":0.201030080408468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1419069490","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31610867,0.0016246059,0.67167956,0.00055537303,0.00012892474,0.00010657701,0.0001763386,0.0016236539,0.007996277],"genre_scores_gemma":[0.94292057,0.0003271467,0.05356952,0.00008415406,0.000024915898,0.000057734025,0.00022210603,0.00005647299,0.002737403],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993975,0.00017058526,0.000037963026,0.0001651119,0.0001619537,0.000066820634],"domain_scores_gemma":[0.9978194,0.0012614225,0.00013073377,0.00014507347,0.0006035828,0.00003984357],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001991248,0.0010842016,0.0010606489,0.0007465453,0.00043522255,0.001363465,0.0013679204,0.0011288583,0.0016555821],"category_scores_gemma":[0.0055311737,0.0003559045,0.00079929334,0.0005956134,0.00030902022,0.0016496998,0.00067343726,0.0011173349,0.00053833623],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003081335,0.00017532973,0.004412302,0.00011938127,0.00010459098,0.00006672369,0.00013804824,0.8307384,0.0024743208,0.0012595126,0.0010170505,0.15918627],"study_design_scores_gemma":[0.0000038225753,0.000026358746,0.00042297042,0.000004392261,0.000010006025,0.000005752079,0.000010760511,0.99866796,0.00048331695,0.0002391835,0.00011970671,0.0000057646],"about_ca_topic_score_codex":0.03641871,"about_ca_topic_score_gemma":0.016213326,"teacher_disagreement_score":0.03641871,"about_ca_system_score_codex":0.001174703,"about_ca_system_score_gemma":0.0011885184,"threshold_uncertainty_score":0.072413445},"labels":[],"label_agreement":null},{"id":"W1480122107","doi":"10.1007/11774129_4","title":"Comparative Analysis of Job Satisfaction in Agile and Non-agile Software Development Teams","year":2006,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":98,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Agile software development; Job satisfaction; Computer science; Software development; Knowledge management; Extreme programming; Job design; Affect (linguistics); Software; Job performance; Software development process; Psychology; Software engineering; Social psychology","score_opus":0.013809750621480347,"score_gpt":0.26204095526622506,"score_spread":0.2482312046447447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1480122107","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99865794,0.00006328548,0.00016933166,0.000032887478,0.000006924466,0.0000054447823,0.00008366209,0.000003705601,0.0009767881],"genre_scores_gemma":[0.9994723,0.000015269743,0.00005023644,0.0000073315628,0.000003543578,0.000006549721,0.00011562985,0.0000021187618,0.00032707205],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99825567,0.0007997815,0.0001179688,0.00009263222,0.0003954685,0.0003385639],"domain_scores_gemma":[0.9831746,0.009367583,0.0017116698,0.00024681079,0.0025301725,0.00296922],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002073044,0.00016423792,0.0003549774,0.0015363423,0.0006238537,0.00085062976,0.0004924531,0.00035396864,0.0038917714],"category_scores_gemma":[0.008309372,0.00012706775,0.00064879237,0.0017668246,0.00038302733,0.0005410353,0.00070166605,0.00040366082,0.0005245766],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003786025,0.0008549744,0.951094,0.00015171008,0.00017183022,0.00027157678,0.0043631704,0.0011770661,0.0018144941,0.00048410837,0.000978068,0.03485289],"study_design_scores_gemma":[0.000029610099,0.0013986696,0.99206173,0.000013064532,0.000033762528,0.000090928916,0.004609517,0.001117919,0.00026110729,0.0000981441,0.00027714635,0.000008390792],"about_ca_topic_score_codex":0.0026559304,"about_ca_topic_score_gemma":0.0031265793,"teacher_disagreement_score":0.0038917714,"about_ca_system_score_codex":0.00055051764,"about_ca_system_score_gemma":0.000604914,"threshold_uncertainty_score":0.013019264},"labels":[],"label_agreement":null},{"id":"W1481457981","doi":"10.1002/spe.2147","title":"Methods for selecting and improving software clustering algorithms","year":2012,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Computer science; Software; Data mining; Algorithm; Strengths and weaknesses; Process (computing); Machine learning; Programming language","score_opus":0.03515311790681467,"score_gpt":0.3848034187302256,"score_spread":0.3496503008234109,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1481457981","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009170644,0.00017066758,0.9975216,0.00009036778,0.000021101734,0.00014193438,0.000020687985,0.0004548875,0.00066175294],"genre_scores_gemma":[0.010134661,0.00018654193,0.98860466,0.000038117698,0.00002372723,0.000276452,0.000087093074,0.00014246076,0.0005063146],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9700382,0.011009334,0.0033584,0.004021549,0.010712932,0.0008596409],"domain_scores_gemma":[0.95659107,0.021275003,0.0042477907,0.0054460694,0.0119809015,0.0004591588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022306519,0.0025136769,0.0018286747,0.008544557,0.0017204719,0.004795381,0.005069713,0.0026870384,0.0037705759],"category_scores_gemma":[0.061016586,0.0016197214,0.0027384835,0.00614393,0.0021414428,0.0053837835,0.0038847625,0.0029482953,0.0020932925],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014193798,0.00018210831,0.003091728,0.001725686,0.00019465301,0.00021848016,0.0015690066,0.039072096,0.012520395,0.10875422,0.00633127,0.82619846],"study_design_scores_gemma":[0.000300747,0.0004081168,0.002685449,0.0011699117,0.00046299814,0.0016455587,0.0013588534,0.6210263,0.06079876,0.1736883,0.13609639,0.00035862954],"about_ca_topic_score_codex":0.0018172467,"about_ca_topic_score_gemma":0.0028926532,"teacher_disagreement_score":0.022306519,"about_ca_system_score_codex":0.0025674247,"about_ca_system_score_gemma":0.0037893527,"threshold_uncertainty_score":0.11796957},"labels":[],"label_agreement":null},{"id":"W1482624438","doi":"","title":"Practical language-independent detection of near-miss clones","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Queen's University","funders":"","keywords":"Computer science; Reuse; Programming language; Code (set theory); Code reuse; Simple (philosophy); Source lines of code; Source code; Software; clone (Java method); Theoretical computer science; Biology; Set (abstract data type)","score_opus":0.01845367035159494,"score_gpt":0.31081303541601557,"score_spread":0.2923593650644206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1482624438","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19263634,0.00048717842,0.777232,0.00059241086,0.000079210644,0.0004922529,0.001128608,0.022480177,0.004871815],"genre_scores_gemma":[0.26649678,0.00016587264,0.72573537,0.00022178564,0.000028766235,0.0002535964,0.0017325162,0.002311474,0.0030538791],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99263316,0.0024165397,0.00091898366,0.0013977605,0.0023063642,0.00032711856],"domain_scores_gemma":[0.9634965,0.017068196,0.002520853,0.007016768,0.009480853,0.0004167883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038634976,0.0011441185,0.001349606,0.0038322797,0.0013850617,0.0022252484,0.0019266595,0.0022490881,0.0056207427],"category_scores_gemma":[0.029126497,0.0008960889,0.00084619113,0.0024697587,0.0010217781,0.003985259,0.003077843,0.0014072828,0.0038589914],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072303344,0.00034144172,0.025505722,0.0018955014,0.00014303024,0.0023203483,0.00457134,0.005165836,0.42494026,0.007467292,0.009301072,0.5176252],"study_design_scores_gemma":[0.000342969,0.0007344933,0.033693034,0.00026552455,0.0003373482,0.010157185,0.0043395073,0.17029643,0.69976056,0.029977826,0.049572956,0.0005222311],"about_ca_topic_score_codex":0.0009270064,"about_ca_topic_score_gemma":0.0017069977,"teacher_disagreement_score":0.0056207427,"about_ca_system_score_codex":0.00056863227,"about_ca_system_score_gemma":0.0010122615,"threshold_uncertainty_score":0.020432353},"labels":[],"label_agreement":null},{"id":"W1483240941","doi":"","title":"Détection de défauts de programmes Java","year":2014,"lang":"fr","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Humanities; Political science; Physics; Philosophy","score_opus":0.020970876611438257,"score_gpt":0.2915658646339321,"score_spread":0.2705949880224938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1483240941","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81727886,0.0010009159,0.14225115,0.0003759749,0.0002148619,0.00017512785,0.0013375176,0.02132084,0.01604477],"genre_scores_gemma":[0.9239881,0.00041305073,0.059652757,0.00022104417,0.00003978888,0.00015859959,0.0016951592,0.0022598794,0.011571697],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99643385,0.00035866242,0.00017231876,0.00047829546,0.0022235182,0.00033329794],"domain_scores_gemma":[0.9878192,0.006117869,0.0016406692,0.0016300235,0.0025030796,0.00028918823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001336975,0.0004952616,0.00031972257,0.0015796294,0.00046225073,0.0010901137,0.00047739947,0.00063619093,0.0020521667],"category_scores_gemma":[0.012232272,0.00033983417,0.0003373718,0.0008344098,0.0006951831,0.0009118241,0.0008963741,0.00089566596,0.00063945586],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021053255,0.0002568914,0.17197031,0.0012066932,0.00012385064,0.002434602,0.007020207,0.009440976,0.33548442,0.018098561,0.011276679,0.44058138],"study_design_scores_gemma":[0.00009258019,0.000726403,0.25289044,0.00046391122,0.00017079478,0.0032947315,0.0012205729,0.14070806,0.506387,0.008600124,0.08525601,0.00018937324],"about_ca_topic_score_codex":0.0036774743,"about_ca_topic_score_gemma":0.0038193003,"teacher_disagreement_score":0.0036774743,"about_ca_system_score_codex":0.00065608515,"about_ca_system_score_gemma":0.000694157,"threshold_uncertainty_score":0.0073121786},"labels":[],"label_agreement":null},{"id":"W1484683231","doi":"10.1007/978-3-540-27769-9_23","title":"Improving Generalization Level in UML Models Iterative Cross Generalization in Practice","year":2004,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Generalization; Unified Modeling Language; Programming language; Theoretical computer science; Algorithm; Artificial intelligence; Mathematics; Software","score_opus":0.03566211871474381,"score_gpt":0.3016579245141278,"score_spread":0.265995805799384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1484683231","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12526019,0.00028564568,0.86249447,0.00041571233,0.000032290034,0.00021048117,0.0000931013,0.006507221,0.004700846],"genre_scores_gemma":[0.5571678,0.00016747702,0.43775344,0.00022933523,0.000028705354,0.00019029086,0.00048722143,0.0013878872,0.0025877862],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9765158,0.010341246,0.0018256123,0.0031888792,0.007234891,0.0008934988],"domain_scores_gemma":[0.91069376,0.043926734,0.0035915233,0.03140169,0.009679383,0.0007069694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016064165,0.0013596656,0.0014507482,0.0024865698,0.0010914021,0.0036751672,0.0029844332,0.0022017392,0.004063083],"category_scores_gemma":[0.07620418,0.0016690412,0.0021328998,0.0015647173,0.0015244328,0.01138602,0.00889555,0.0036894074,0.0013065179],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006633575,0.0010604493,0.022347534,0.00068180956,0.00037981782,0.0004602928,0.0071842573,0.10892205,0.035280317,0.061327253,0.0044359276,0.7572568],"study_design_scores_gemma":[0.00006429963,0.0003854069,0.0032182068,0.00021415373,0.0003538899,0.00039049753,0.0007469693,0.87935144,0.046817325,0.058035832,0.010345398,0.00007659561],"about_ca_topic_score_codex":0.0032703078,"about_ca_topic_score_gemma":0.00463256,"teacher_disagreement_score":0.016064165,"about_ca_system_score_codex":0.001936975,"about_ca_system_score_gemma":0.0024660174,"threshold_uncertainty_score":0.08495647},"labels":[],"label_agreement":null},{"id":"W1485562483","doi":"10.1007/978-3-540-87879-7_3","title":"A Tool to Visualize Architectural Design Decisions","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Visualization; Decision analysis; Software engineering; Decision support system; Architecture; Software architecture; Business decision mapping; Key (lock); Data science; Software; Management science; Artificial intelligence; Engineering; Programming language","score_opus":0.0410595540471483,"score_gpt":0.29701895047819976,"score_spread":0.2559593964310515,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1485562483","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010140162,0.0005908102,0.8700344,0.00050090055,0.00024030298,0.000193886,0.0065850113,0.08756202,0.024152417],"genre_scores_gemma":[0.05282057,0.0013128138,0.90307117,0.00023881889,0.0000736864,0.00048891775,0.011452423,0.010801024,0.019740654],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996388,0.00008076088,0.00004133376,0.000051845534,0.00014933942,0.000037954593],"domain_scores_gemma":[0.99796486,0.0012035357,0.00014907669,0.00030625446,0.0002523405,0.00012397647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008841374,0.0024750328,0.00079376955,0.004104827,0.0008890688,0.003819747,0.0013334126,0.0016098643,0.034875773],"category_scores_gemma":[0.0034441852,0.0013062367,0.0012880057,0.0026388967,0.00046199342,0.00332747,0.002006815,0.0023692506,0.008096318],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035559724,0.00037831042,0.003662134,0.0015823301,0.00011922562,0.001107332,0.0028464727,0.03449981,0.020104572,0.06888282,0.21661304,0.64984834],"study_design_scores_gemma":[0.00024421557,0.00016295894,0.002498014,0.0009352552,0.00017990565,0.0012764832,0.00072880444,0.35172802,0.02048135,0.08059317,0.5409388,0.00023300044],"about_ca_topic_score_codex":0.0037474763,"about_ca_topic_score_gemma":0.008680814,"teacher_disagreement_score":0.034875773,"about_ca_system_score_codex":0.0005500058,"about_ca_system_score_gemma":0.0010165912,"threshold_uncertainty_score":0.116671085},"labels":[],"label_agreement":null},{"id":"W1485825596","doi":"","title":"Non-Functional Requirements: Size Measurement and Testing with COSMIC-FFP","year":2007,"lang":"en","type":"article","venue":"University of Twente Research Information","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Reliability engineering; Process (computing); Non-functional requirement; Systems engineering; Scope (computer science); Functional requirement; Functional testing; Software engineering; Software; Engineering; Software development; Software construction","score_opus":0.07659704305500299,"score_gpt":0.268884615089515,"score_spread":0.19228757203451202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1485825596","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07169458,0.0001377798,0.9103599,0.00031163308,0.000026939968,0.00028158107,0.00059474545,0.0046513923,0.0119413845],"genre_scores_gemma":[0.49186534,0.00008794797,0.50547314,0.00008225336,0.00001368457,0.0003872319,0.00086670555,0.0003209321,0.0009028786],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98656565,0.0059566693,0.00068525545,0.00072952814,0.005727664,0.0003352466],"domain_scores_gemma":[0.9638103,0.01856386,0.0024341992,0.008412837,0.0064343354,0.00034440137],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062462236,0.0011156598,0.00056786695,0.004473597,0.0005568811,0.0018151203,0.001956339,0.0013063413,0.0021389665],"category_scores_gemma":[0.038693532,0.00048099822,0.000883574,0.0027843984,0.0013862038,0.0023661566,0.0014594696,0.0009402014,0.00045763687],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005297724,0.0003317274,0.03407922,0.0005944293,0.00017284698,0.0011550528,0.0013972416,0.23629208,0.027024643,0.11446208,0.009986726,0.57397425],"study_design_scores_gemma":[0.000049595288,0.00025685778,0.01511566,0.00012020325,0.000068009664,0.00090605207,0.0002654155,0.9100128,0.035400122,0.029445816,0.008267447,0.00009201891],"about_ca_topic_score_codex":0.006796109,"about_ca_topic_score_gemma":0.004835636,"teacher_disagreement_score":0.006796109,"about_ca_system_score_codex":0.0015692262,"about_ca_system_score_gemma":0.0014621593,"threshold_uncertainty_score":0.03303361},"labels":[],"label_agreement":null},{"id":"W1486880988","doi":"","title":"A Service Sharing Approach to Integrating Program Comprehension Tools","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Program comprehension; Computer science; Software engineering; Software maintenance; Comprehension; Flexibility (engineering); Software development; Software; Systems development life cycle; Software development process; Software system; Operating system","score_opus":0.06299873924088611,"score_gpt":0.30341982281097635,"score_spread":0.24042108357009023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1486880988","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005983916,0.00007883289,0.98053944,0.00073475845,0.000048363792,0.00026573744,0.00003821154,0.0035665615,0.008744207],"genre_scores_gemma":[0.13846706,0.00029844127,0.8452064,0.0006156067,0.00011694306,0.0006960637,0.0003482895,0.0008936216,0.013357599],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99309915,0.002430873,0.00055729435,0.0009154228,0.002300654,0.0006967445],"domain_scores_gemma":[0.9888855,0.002189912,0.0007557131,0.00516678,0.0021921634,0.00080983527],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007064431,0.0014403606,0.0009768584,0.0033561827,0.0027303125,0.004544798,0.00512422,0.0035479362,0.00643387],"category_scores_gemma":[0.014389361,0.0011521501,0.0020725974,0.0032847389,0.0027305877,0.011947581,0.009340853,0.003449743,0.003280868],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036911012,0.0010849702,0.005527984,0.0005082741,0.00020196212,0.0018510786,0.012289159,0.012483853,0.03213798,0.41818944,0.013210043,0.5021462],"study_design_scores_gemma":[0.00016432567,0.0009261695,0.0024934046,0.0003624574,0.0005747108,0.004087167,0.0033257077,0.18176642,0.04330797,0.36211696,0.40055946,0.00031528287],"about_ca_topic_score_codex":0.005690308,"about_ca_topic_score_gemma":0.0045916485,"teacher_disagreement_score":0.007064431,"about_ca_system_score_codex":0.0019653516,"about_ca_system_score_gemma":0.0047095725,"threshold_uncertainty_score":0.037360728},"labels":[],"label_agreement":null},{"id":"W1488045480","doi":"10.1109/wcre.2012.53","title":"Analyzing the Impact of Antipatterns on Change-Proneness Using Fine-Grained Source Code Changes","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Java; Source code; Open source; Computer science; Biology; Programming language; Software","score_opus":0.1016252506006192,"score_gpt":0.35390155395753703,"score_spread":0.25227630335691786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1488045480","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9927199,0.00017250396,0.0058681504,0.000026290609,0.000009822978,0.000051740084,0.00055389374,0.00015281765,0.0004449833],"genre_scores_gemma":[0.9946051,0.000059162754,0.004274537,0.000009573641,0.000008546359,0.000023756125,0.0008181592,0.000034706696,0.00016646112],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9951794,0.0008320527,0.0006740379,0.0012461548,0.0018202538,0.00024811891],"domain_scores_gemma":[0.8859812,0.06887137,0.02575161,0.010917935,0.0073055048,0.0011723564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026983952,0.0004046421,0.00034207423,0.0047703977,0.00030488815,0.0006943343,0.00036535988,0.00036060292,0.00076263107],"category_scores_gemma":[0.030424414,0.0002434549,0.00049039343,0.0025458618,0.00055846607,0.0010583072,0.0007349388,0.00060463184,0.00013945019],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028427635,0.00012661534,0.9156278,0.00020975221,0.00026988974,0.0002616078,0.00058946916,0.0074514793,0.015325792,0.00020024108,0.00020298437,0.05945006],"study_design_scores_gemma":[0.0000056144713,0.00015328458,0.98373556,0.000013107005,0.000069707006,0.00021987197,0.00015348586,0.010974283,0.0037515564,0.00026336042,0.00064421893,0.000016020917],"about_ca_topic_score_codex":0.0019690262,"about_ca_topic_score_gemma":0.0038296282,"teacher_disagreement_score":0.0047703977,"about_ca_system_score_codex":0.00035723398,"about_ca_system_score_gemma":0.00035892442,"threshold_uncertainty_score":0.014270604},"labels":[],"label_agreement":null},{"id":"W1489439366","doi":"10.1016/b978-0-12-411519-4.00019-7","title":"Analytical Product Release Planning","year":2015,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Product (mathematics); Business; Computer science; Process management; Process engineering; Mathematics; Engineering; Geometry","score_opus":0.04954594433010383,"score_gpt":0.29727474127972153,"score_spread":0.2477287969496177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1489439366","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004971564,0.010206246,0.6598525,0.0011099238,0.0005553725,0.00018721909,0.0007500443,0.0028992847,0.31946784],"genre_scores_gemma":[0.16196656,0.015854465,0.3936032,0.00033540375,0.00042114698,0.0003016707,0.002825521,0.0016678686,0.42302424],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991598,0.00013806773,0.000037298927,0.00013978331,0.00045780715,0.000067257264],"domain_scores_gemma":[0.99911875,0.00038586478,0.000052154155,0.0001975442,0.00020088234,0.000044802484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011574853,0.0016053559,0.00079757505,0.002204387,0.00068000815,0.0032948237,0.0014069182,0.00071485154,0.036056038],"category_scores_gemma":[0.00260454,0.0011089736,0.00092820224,0.0025205957,0.0006522118,0.0019357584,0.0011162886,0.0013767863,0.0133897625],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000067738256,0.00006574145,0.00031066133,0.00026863196,0.00003147449,0.00008862683,0.00011178688,0.07220462,0.0024715594,0.07064374,0.05701334,0.7967222],"study_design_scores_gemma":[0.000025639789,0.00013946673,0.0010899986,0.0004936595,0.000061166975,0.00033275178,0.00022665174,0.33448908,0.00702748,0.22992562,0.4261133,0.00007523615],"about_ca_topic_score_codex":0.0042267763,"about_ca_topic_score_gemma":0.0060019707,"teacher_disagreement_score":0.036056038,"about_ca_system_score_codex":0.001342776,"about_ca_system_score_gemma":0.0020043221,"threshold_uncertainty_score":0.120619476},"labels":[],"label_agreement":null},{"id":"W1490878083","doi":"10.1109/wpc.1999.777743","title":"Extending software quality assessment techniques to Java systems","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Java; Computer science; Software quality; Source code; Suite; Software system; Real time Java; Java annotation; Software engineering; Programming language; Generics in Java; Software; Operating system; Software development","score_opus":0.051781090674672445,"score_gpt":0.36953608612818306,"score_spread":0.3177549954535106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1490878083","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0075746356,0.0005033107,0.986956,0.00018277552,0.000030222263,0.00017487178,0.000048259728,0.0019346983,0.0025951762],"genre_scores_gemma":[0.13864377,0.0008854944,0.8577749,0.00013682478,0.00006337151,0.00027571837,0.00022272335,0.00036117109,0.0016360937],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98460776,0.0034603183,0.0015239643,0.00095733296,0.009094129,0.00035648394],"domain_scores_gemma":[0.9642854,0.015325316,0.0033383684,0.0040842243,0.012580548,0.0003861988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00842067,0.0010459876,0.0009436601,0.005884847,0.0005870851,0.0025296658,0.0013149475,0.000583767,0.0012108069],"category_scores_gemma":[0.04372933,0.00043335735,0.0009796489,0.0049787555,0.0009014618,0.0026984198,0.0020537758,0.0019812954,0.00061150984],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000043340904,0.00012486783,0.006573722,0.0005286342,0.000091121365,0.00016309864,0.00096502574,0.016818643,0.009758784,0.01632956,0.0017483444,0.9468549],"study_design_scores_gemma":[0.00014554513,0.00090695807,0.031426672,0.0015901526,0.0004905754,0.0023818614,0.0010861795,0.586671,0.053408258,0.17639522,0.14514019,0.00035735487],"about_ca_topic_score_codex":0.0055817696,"about_ca_topic_score_gemma":0.006042382,"teacher_disagreement_score":0.00842067,"about_ca_system_score_codex":0.001074232,"about_ca_system_score_gemma":0.0016167022,"threshold_uncertainty_score":0.044533312},"labels":[],"label_agreement":null},{"id":"W1491093693","doi":"10.1007/978-3-642-13428-9_5","title":"Empirical Evaluation of Selected Algorithms for Complexity-Based Classification of Software Modules and a New Model","year":2010,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Machine learning; Computer science; Artificial intelligence; Software quality; Data mining; Support vector machine; Software; Classifier (UML); Linear discriminant analysis; Software metric; Naive Bayes classifier; Categorization; Statistical classification; Algorithm; Software development","score_opus":0.4482315763401465,"score_gpt":0.4587910225596661,"score_spread":0.010559446219519597,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1491093693","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84205157,0.0036557221,0.14660852,0.0007230413,0.00015731606,0.0003625962,0.0014609694,0.0010130664,0.003967242],"genre_scores_gemma":[0.8829701,0.00063651643,0.11098481,0.00007423364,0.000083286046,0.00024355829,0.0038825122,0.00019993426,0.00092516316],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9883591,0.0064471443,0.0011635451,0.0013520252,0.002325891,0.0003522874],"domain_scores_gemma":[0.73198897,0.24324732,0.004260239,0.010100901,0.008929883,0.0014727789],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020833448,0.0015841383,0.0013070444,0.005939205,0.0008521586,0.0026825531,0.0036506266,0.0024226557,0.002004101],"category_scores_gemma":[0.09557428,0.0003731699,0.0014019888,0.00433629,0.001461158,0.0049712965,0.0022803717,0.0018830587,0.0006184575],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007427089,0.0025275412,0.15423238,0.0012183446,0.0014357845,0.000098601704,0.0011254072,0.2053584,0.001824458,0.009530048,0.010714972,0.60450697],"study_design_scores_gemma":[0.00033997826,0.0011840677,0.03428017,0.000091298716,0.00036374887,0.0002124168,0.00049629394,0.9506857,0.0025334451,0.008774553,0.0009651885,0.00007308983],"about_ca_topic_score_codex":0.005325634,"about_ca_topic_score_gemma":0.0043776887,"teacher_disagreement_score":0.020833448,"about_ca_system_score_codex":0.0029343045,"about_ca_system_score_gemma":0.0014961011,"threshold_uncertainty_score":0.110179126},"labels":[],"label_agreement":null},{"id":"W1492274812","doi":"10.1002/smr.1644","title":"Model refactoring using examples: a search‐based approach","year":2014,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Code refactoring; Computer science; Set (abstract data type); Heuristic; Metamodeling; Quality (philosophy); Base (topology); Sequence (biology); Software; Programming language; Software engineering; Artificial intelligence","score_opus":0.06070729225650119,"score_gpt":0.3009794474827227,"score_spread":0.24027215522622153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1492274812","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027760778,0.00047820908,0.96601516,0.0002837078,0.000018449768,0.00022385996,0.000080628015,0.001055867,0.004083394],"genre_scores_gemma":[0.14624862,0.00027083178,0.85182834,0.000069064445,0.000014776798,0.00021365142,0.00025417798,0.00016888765,0.00093169766],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99730635,0.0014845724,0.00016516382,0.00031361362,0.0006179863,0.000112288726],"domain_scores_gemma":[0.99281704,0.005346529,0.0004245496,0.0005911784,0.00069643965,0.00012433823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027081475,0.0014128933,0.0011941412,0.0035619861,0.0006644027,0.001557912,0.0027573602,0.0018700643,0.0032289082],"category_scores_gemma":[0.010137625,0.0008024334,0.0013909314,0.0022260617,0.0013364849,0.0014914356,0.0018366982,0.0010639695,0.00052666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020948451,0.00040485812,0.0028167951,0.00056978397,0.0001850866,0.00046474946,0.0004683066,0.6598517,0.0067160656,0.040055223,0.002306557,0.2859514],"study_design_scores_gemma":[0.000041698906,0.00007082144,0.00021831002,0.00007248135,0.000046658086,0.00011781232,0.00008033873,0.98359144,0.0025973993,0.01031559,0.0028297186,0.000017659793],"about_ca_topic_score_codex":0.0030181631,"about_ca_topic_score_gemma":0.0039609326,"teacher_disagreement_score":0.0035619861,"about_ca_system_score_codex":0.00080803817,"about_ca_system_score_gemma":0.0010613626,"threshold_uncertainty_score":0.014322221},"labels":[],"label_agreement":null},{"id":"W1493296086","doi":"10.1007/s10664-014-9308-x","title":"Understanding the impact of rapid releases on software quality","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal","funders":"","keywords":"Computer science; Software release life cycle; Software engineering; Crash; Software quality; Software; Quality (philosophy); Software bug; Software quality analyst; Software quality assurance; Software development; Operating system","score_opus":0.11284751522180525,"score_gpt":0.3485465269864667,"score_spread":0.23569901176466146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1493296086","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9870481,0.0006787787,0.0019905353,0.0014675538,0.00001750796,0.000012565096,0.00010719434,0.000041481962,0.008636277],"genre_scores_gemma":[0.9986981,0.00021926571,0.00034957053,0.00006432055,0.000019238118,0.0000036127674,0.000048109217,0.000015984699,0.00058164774],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9958625,0.0018513605,0.00016376864,0.00035758744,0.0012474306,0.00051732967],"domain_scores_gemma":[0.8305765,0.13563432,0.020508679,0.0040833363,0.006711033,0.002486182],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006575145,0.00033545005,0.00022169376,0.0012466682,0.0003639849,0.002650788,0.00063172734,0.0007225048,0.0053715413],"category_scores_gemma":[0.082188606,0.0003142977,0.00034340506,0.0010755843,0.0009780439,0.0045724185,0.0010129744,0.0017584773,0.00046099882],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012849903,0.0020655645,0.73646396,0.00060333934,0.0004948434,0.0006678892,0.0035674218,0.030894596,0.0095191,0.030389568,0.003547351,0.18050137],"study_design_scores_gemma":[0.00008966751,0.0009472173,0.942774,0.000122031444,0.00025155288,0.00013419327,0.003019379,0.028830081,0.0027144544,0.01775897,0.0033086205,0.000049856204],"about_ca_topic_score_codex":0.007476118,"about_ca_topic_score_gemma":0.009481686,"teacher_disagreement_score":0.007476118,"about_ca_system_score_codex":0.0014474343,"about_ca_system_score_gemma":0.0013111045,"threshold_uncertainty_score":0.03477317},"labels":[],"label_agreement":null},{"id":"W1493964280","doi":"","title":"Proceedings of the 2008 international working conference on Mining software repositories","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Waterloo; Queen's University","funders":"","keywords":"Schedule; Scope (computer science); Maturity (psychological); Library science; Competition (biology); Operations research; Computer science; Political science; Public relations; Data science; Engineering; Law","score_opus":0.039903117125145875,"score_gpt":0.2552659880458374,"score_spread":0.2153628709206915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1493964280","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043395646,0.09287303,0.48260382,0.105691195,0.056271553,0.00414191,0.034594405,0.028262548,0.15216593],"genre_scores_gemma":[0.058686417,0.040985513,0.3626462,0.009078234,0.008990698,0.002543828,0.11954407,0.008372175,0.38915285],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9883659,0.004103254,0.0010318728,0.0021622642,0.0037021535,0.00063468015],"domain_scores_gemma":[0.9721432,0.008447736,0.0014101113,0.0058098105,0.008037175,0.0041518863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016700944,0.0017481926,0.0024461243,0.0074573103,0.0019687272,0.013420267,0.0032608188,0.0029809636,0.07071235],"category_scores_gemma":[0.036420517,0.0013254624,0.0027828761,0.006007185,0.0019208586,0.013252309,0.005660405,0.00508519,0.042674445],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020773626,0.00036495415,0.0029576526,0.00080569304,0.00022009759,0.00017992037,0.00075874384,0.0009677366,0.0020695026,0.0061869165,0.6113573,0.37392372],"study_design_scores_gemma":[0.00006232919,0.00022015099,0.006242973,0.0010422638,0.0001366375,0.00048044507,0.00096996594,0.0052414765,0.0021481046,0.015709652,0.96763736,0.00010864371],"about_ca_topic_score_codex":0.0046943445,"about_ca_topic_score_gemma":0.0080420235,"teacher_disagreement_score":0.07071235,"about_ca_system_score_codex":0.002371779,"about_ca_system_score_gemma":0.006435498,"threshold_uncertainty_score":0.23655641},"labels":[],"label_agreement":null},{"id":"W1494567756","doi":"","title":"Program comprehension with dynamic recovery of code collaboration patterns and roles","year":2004,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Legacy system; COBOL; Program comprehension; Computer science; Maintainability; Software maintenance; Software engineering; Legacy code; Artifact (error); Source code; Software development; Software system; Software quality; Code (set theory); Programming language; Software; Artificial intelligence","score_opus":0.04702776759363301,"score_gpt":0.3754114805347369,"score_spread":0.32838371294110386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1494567756","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014975389,0.000057003366,0.97758406,0.00030042403,0.000013005594,0.00012926119,0.00009565597,0.00585777,0.0009873648],"genre_scores_gemma":[0.1177738,0.00014488804,0.87753934,0.00010315118,0.000024836685,0.0002253534,0.0006860102,0.00097471016,0.002527848],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959876,0.001155635,0.00025468308,0.0011140588,0.0012460023,0.00024203313],"domain_scores_gemma":[0.9886952,0.0044295765,0.0016700882,0.0035406447,0.0014768121,0.00018766358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028379397,0.0013514542,0.0009930151,0.0036118669,0.0013619758,0.002908214,0.0028082042,0.001705065,0.0024929694],"category_scores_gemma":[0.015841078,0.0008505762,0.0014508498,0.002046273,0.0024069233,0.0054576984,0.004332187,0.002855213,0.0011307287],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028388452,0.00046415784,0.0064109066,0.0007279874,0.00009708738,0.0009377482,0.011057559,0.021419749,0.07772608,0.057579894,0.0070373714,0.81625766],"study_design_scores_gemma":[0.00014243959,0.00033517438,0.0065756617,0.000272138,0.000211082,0.0029533755,0.0037545788,0.6057923,0.15665655,0.13995756,0.08312194,0.00022724301],"about_ca_topic_score_codex":0.0019720946,"about_ca_topic_score_gemma":0.0021919757,"teacher_disagreement_score":0.0036118669,"about_ca_system_score_codex":0.0008431299,"about_ca_system_score_gemma":0.0024431183,"threshold_uncertainty_score":0.015008628},"labels":[],"label_agreement":null},{"id":"W1494902786","doi":"10.1007/978-0-387-21599-0_13","title":"The SPOOL Design Repository: Architecture, Schema, and Mechanisms","year":2002,"lang":"en","type":"book-chapter","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bell (Canada); Université de Montréal","funders":"","keywords":"Reverse engineering; Suite; Software engineering; Computer science; Schema (genetic algorithms); Process (computing); Systems engineering; Engineering; Programming language","score_opus":0.010191757098158567,"score_gpt":0.21246413789952628,"score_spread":0.20227238080136772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1494902786","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004698443,0.009744513,0.8760007,0.005064449,0.0008057003,0.00027457706,0.0023793732,0.017926667,0.08310552],"genre_scores_gemma":[0.05368979,0.017878067,0.7409863,0.0025701,0.0005501201,0.0006385419,0.011904815,0.010403041,0.16137926],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970902,0.0006858368,0.00032955868,0.0003463332,0.0013747925,0.00017321813],"domain_scores_gemma":[0.99447834,0.0013282857,0.0002350902,0.00251062,0.0010419812,0.00040566444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054452317,0.00094694935,0.0014013542,0.0047707553,0.0017357458,0.013733463,0.0042948425,0.002271441,0.01611564],"category_scores_gemma":[0.010054381,0.0019583874,0.0010438076,0.008034776,0.0032787588,0.017041715,0.004600095,0.004439163,0.015070116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000871167,0.00009794738,0.0010432402,0.00049672194,0.000042826458,0.00016352165,0.0015127624,0.0025290332,0.0018752385,0.4029519,0.13713203,0.45206764],"study_design_scores_gemma":[0.00003693673,0.000035103614,0.0002552059,0.0004714536,0.000039297087,0.000729317,0.0002847318,0.007632658,0.0027502293,0.16931593,0.81838316,0.000065973116],"about_ca_topic_score_codex":0.0027059861,"about_ca_topic_score_gemma":0.0038889754,"teacher_disagreement_score":0.01611564,"about_ca_system_score_codex":0.0021313508,"about_ca_system_score_gemma":0.0057018967,"threshold_uncertainty_score":0.053912163},"labels":[],"label_agreement":null},{"id":"W1496838008","doi":"10.1109/iwpse.2004.11","title":"Evolution Spectrographs: visualizing punctuated change in software evolution","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Punctuated equilibrium; Computer science; Software; Software visualization; Punctuation; Visualization; Spectrograph; Process (computing); Software development; Data mining; Artificial intelligence; Software construction; Paleontology; Geology; Programming language; Astronomy","score_opus":0.025745816601175902,"score_gpt":0.28801124415715257,"score_spread":0.26226542755597665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1496838008","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4656933,0.0013820238,0.5000803,0.0009991784,0.0001521462,0.00024229563,0.004153264,0.017330628,0.009966888],"genre_scores_gemma":[0.805203,0.000623204,0.18965508,0.0001312464,0.000055127603,0.00013416675,0.001510959,0.00064115756,0.0020461937],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999746,0.00007750143,0.000015912661,0.000042462336,0.00008680079,0.000031269778],"domain_scores_gemma":[0.9977708,0.001137441,0.00042047352,0.00016793543,0.00034761673,0.00015571475],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007284467,0.00040662853,0.00017820892,0.0050824066,0.00037554922,0.0007932259,0.00039455362,0.00069659005,0.002299247],"category_scores_gemma":[0.0039858934,0.00022039528,0.00026269702,0.002393824,0.0003331063,0.0014386066,0.0010008968,0.00076030433,0.00026457667],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012051389,0.00027167227,0.06877828,0.0009800064,0.0002886343,0.0019016347,0.016706409,0.024076702,0.18176962,0.022167234,0.02488526,0.6569694],"study_design_scores_gemma":[0.00017161897,0.0005005247,0.33569947,0.00044884806,0.0002480121,0.0037698073,0.0054286644,0.40172932,0.10951897,0.03503426,0.10704898,0.00040147876],"about_ca_topic_score_codex":0.0039991704,"about_ca_topic_score_gemma":0.0042737825,"teacher_disagreement_score":0.0050824066,"about_ca_system_score_codex":0.0003979634,"about_ca_system_score_gemma":0.00039725247,"threshold_uncertainty_score":0.007951796},"labels":[],"label_agreement":null},{"id":"W1498869709","doi":"10.1109/step.2004.14","title":"Predictive Software Models","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software engineering; Software; Software construction; Software sizing; Verification and validation; Software metric; Predictive modelling; Software development; Data mining; Machine learning; Engineering; Operating system","score_opus":0.013616098987263407,"score_gpt":0.2297650146816346,"score_spread":0.2161489156943712,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1498869709","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030442234,0.001070247,0.9351087,0.0016409154,0.00019040462,0.0002810136,0.015029889,0.0045439303,0.011692659],"genre_scores_gemma":[0.6263431,0.0026172458,0.31927145,0.0007048775,0.00033174502,0.0013973482,0.036976974,0.00095448946,0.011402903],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9970251,0.00079844607,0.00016781551,0.00091718236,0.00091463316,0.0001768993],"domain_scores_gemma":[0.9845354,0.009861293,0.0011362551,0.0027230736,0.0015626026,0.00018140454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036503393,0.001558122,0.0011356458,0.0036772508,0.0006487611,0.0029529708,0.003506643,0.0016289295,0.008415914],"category_scores_gemma":[0.026544316,0.00085806136,0.0019290872,0.0038826815,0.0009621502,0.005459308,0.0014742146,0.0027933724,0.0034650203],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016780681,0.00017769221,0.0109941885,0.00039891325,0.00025343779,0.00012882451,0.00014785842,0.7806959,0.00086914474,0.08417433,0.019879455,0.102112494],"study_design_scores_gemma":[0.000019261457,0.000029588435,0.0009012499,0.00005057654,0.000043500353,0.00004970408,0.000027098014,0.918941,0.00047554736,0.07297716,0.0064633586,0.000021913356],"about_ca_topic_score_codex":0.008876162,"about_ca_topic_score_gemma":0.010626225,"teacher_disagreement_score":0.008876162,"about_ca_system_score_codex":0.0016637002,"about_ca_system_score_gemma":0.0019540128,"threshold_uncertainty_score":0.028154016},"labels":[],"label_agreement":null},{"id":"W1498914194","doi":"10.1007/978-3-540-45099-3_11","title":"Efficient Inference of Static Types for Java Bytecode","year":2000,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Bytecode; Java bytecode; Computer science; Type inference; Programming language; Java; Static analysis; Inference; Representation (politics); Program analysis; Theoretical computer science; Algorithm; Java applet; Artificial intelligence; Java annotation","score_opus":0.01949647408175968,"score_gpt":0.276019564644905,"score_spread":0.2565230905631453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1498914194","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009956432,0.000354088,0.9641845,0.00024015384,0.00013523427,0.00008763531,0.00059694366,0.019719236,0.0047256257],"genre_scores_gemma":[0.18431868,0.0006235086,0.7910711,0.00039023045,0.00028130977,0.00018133405,0.0036886495,0.0076026213,0.011842494],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9957775,0.0008465321,0.00033969412,0.00067273487,0.0018745376,0.00048903615],"domain_scores_gemma":[0.9891861,0.0065047625,0.00038600483,0.002501507,0.0012602508,0.00016146807],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025199482,0.0014100792,0.0020686847,0.0028489195,0.0014311321,0.0036700235,0.0055115055,0.0018097104,0.012142959],"category_scores_gemma":[0.014354229,0.0026677244,0.003755544,0.0027149068,0.0018358782,0.0074188407,0.0040164595,0.003350121,0.006460281],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006168266,0.0002706895,0.0031826652,0.0008064112,0.00020364462,0.0003450636,0.00054388883,0.027196985,0.021956317,0.111064754,0.03036984,0.8034429],"study_design_scores_gemma":[0.0001704223,0.00009970703,0.0019323258,0.00020070162,0.00029214835,0.00041259153,0.0002116018,0.4701194,0.056748305,0.44191265,0.02776211,0.0001380689],"about_ca_topic_score_codex":0.005021555,"about_ca_topic_score_gemma":0.009361416,"teacher_disagreement_score":0.012142959,"about_ca_system_score_codex":0.0019320971,"about_ca_system_score_gemma":0.002513079,"threshold_uncertainty_score":0.040622234},"labels":[],"label_agreement":null},{"id":"W1499163112","doi":"10.1007/978-3-642-17578-7_9","title":"Exploring Empirically the Relationship between Lack of Cohesion and Testability in Object-Oriented Systems","year":2010,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Testability; Cohesion (chemistry); Computer science; Object-oriented programming; Engineering; Reliability engineering; Programming language; Chemistry","score_opus":0.28110740335610207,"score_gpt":0.36295490499261546,"score_spread":0.08184750163651339,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1499163112","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98980147,0.00027463638,0.0059186034,0.00035352414,0.0000060256316,0.000014231996,0.000018722578,0.000013165445,0.0035996858],"genre_scores_gemma":[0.998049,0.00008176791,0.0015552535,0.000015238433,0.000008839183,0.000013506399,0.000029447734,0.000007182615,0.0002398793],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99690396,0.001778527,0.00015485736,0.00020399537,0.0007929563,0.0001657307],"domain_scores_gemma":[0.6629754,0.318651,0.012274897,0.002643877,0.0025557268,0.00089913676],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045100986,0.00035492616,0.00029618153,0.0013973142,0.0005144348,0.0017618773,0.0007095906,0.0009712184,0.002087416],"category_scores_gemma":[0.113073885,0.00048220163,0.00034638913,0.0016195286,0.002021145,0.003586592,0.0011560671,0.0016651312,0.00010105792],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037506936,0.0006935565,0.8729691,0.00028049794,0.00027969445,0.00042688716,0.007470147,0.013004216,0.0035084845,0.032765288,0.0004923423,0.067734726],"study_design_scores_gemma":[0.000044145145,0.000788744,0.8693024,0.00010162649,0.00021750627,0.000514948,0.004960379,0.04877518,0.0022663472,0.07213917,0.00084286724,0.000046555917],"about_ca_topic_score_codex":0.0022087346,"about_ca_topic_score_gemma":0.0030923912,"teacher_disagreement_score":0.0045100986,"about_ca_system_score_codex":0.00069161004,"about_ca_system_score_gemma":0.000728739,"threshold_uncertainty_score":0.02385199},"labels":[],"label_agreement":null},{"id":"W1499997010","doi":"","title":"Removing false code dependencies to speedup software build processes","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); University of Toronto","funders":"","keywords":"Computer science; Preprocessor; Header; Speedup; Source code; Software; Programming language; Code (set theory); Graph; Static program analysis; Parallel computing; Software development; Theoretical computer science; Set (abstract data type)","score_opus":0.024592168633208955,"score_gpt":0.2754633940127593,"score_spread":0.25087122537955037,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1499997010","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12813534,0.00037488458,0.8436357,0.00071429927,0.00010808687,0.00011364986,0.00017677213,0.024849115,0.0018920973],"genre_scores_gemma":[0.25338486,0.00026617103,0.73862237,0.0002792155,0.000055732493,0.000102131766,0.0009371016,0.003569037,0.00278337],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976426,0.0005433663,0.00013423798,0.00034832,0.0011190249,0.00021241671],"domain_scores_gemma":[0.9712384,0.016488476,0.0019950583,0.007204718,0.0026324124,0.00044101992],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001507348,0.0015245634,0.00085304776,0.0019360082,0.00096020696,0.0014315944,0.0025083022,0.0016018815,0.0039121527],"category_scores_gemma":[0.017851839,0.0009932758,0.000926666,0.0018718471,0.0010770259,0.0039548026,0.0017016926,0.002615613,0.0021204452],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062150205,0.0008097774,0.010579036,0.00069911376,0.00015348209,0.0008737639,0.0009323807,0.097008705,0.19946446,0.01978865,0.016840272,0.6522287],"study_design_scores_gemma":[0.00014987765,0.0004445203,0.005916932,0.000070685004,0.0002276445,0.0009727196,0.00019122062,0.70220876,0.23602359,0.037008457,0.01667916,0.000106443695],"about_ca_topic_score_codex":0.001947111,"about_ca_topic_score_gemma":0.004042896,"teacher_disagreement_score":0.0039121527,"about_ca_system_score_codex":0.00060062966,"about_ca_system_score_gemma":0.0019058272,"threshold_uncertainty_score":0.013087451},"labels":[],"label_agreement":null},{"id":"W1500112013","doi":"10.4018/978-1-930708-41-9.ch019","title":"The Role of Use Cases in the UML","year":2002,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland; University of Lethbridge","funders":"","keywords":"Computer science; Class diagram; Class (philosophy); Use Case Diagram; Variety (cybernetics); Unified Modeling Language; Categorization; Abstraction; Database transaction; Software engineering; Programming language; Software; Artificial intelligence; Epistemology","score_opus":0.028584457685930154,"score_gpt":0.2518700343251749,"score_spread":0.22328557663924475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1500112013","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013141982,0.0051543824,0.861328,0.013060806,0.00056453975,0.00068863557,0.00027001972,0.0012215482,0.10457003],"genre_scores_gemma":[0.35691085,0.006623571,0.6170113,0.0017271162,0.00057354156,0.0014822228,0.0005727037,0.00064122403,0.014457471],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9046026,0.069853604,0.0060765096,0.004255337,0.0131899845,0.0020219125],"domain_scores_gemma":[0.851187,0.116425745,0.006914298,0.014490633,0.009210865,0.0017714305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05292249,0.0013916821,0.0012590543,0.00938694,0.004930828,0.024048783,0.0043880427,0.0050405,0.0039895866],"category_scores_gemma":[0.08087819,0.0021064167,0.0014809594,0.008161142,0.025918374,0.043605432,0.0068727382,0.0063986015,0.0012991398],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008409562,0.0000118091275,0.00046954583,0.00006898448,0.000006294334,0.00015030439,0.006941348,0.0006458556,0.00013117022,0.97329766,0.0011183793,0.017150281],"study_design_scores_gemma":[0.000017308557,0.000030899835,0.000457027,0.00068095746,0.000023990615,0.00063161156,0.0039301915,0.008364614,0.0006458657,0.7419389,0.24321274,0.00006593819],"about_ca_topic_score_codex":0.008509219,"about_ca_topic_score_gemma":0.0044836635,"teacher_disagreement_score":0.05292249,"about_ca_system_score_codex":0.00854644,"about_ca_system_score_gemma":0.007621199,"threshold_uncertainty_score":0.27988422},"labels":[],"label_agreement":null},{"id":"W1500946169","doi":"10.1145/1368088.1368123","title":"TODO or to bug","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":124,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; University of Victoria","funders":"","keywords":"Computer science; Software development; Task (project management); Software engineering; Team software process; Task management; Variety (cybernetics); Process (computing); Software construction; Personal software process; Software development process; Software; Open-source software development; Empirical research; Software maintenance; Source code; Programming language; Artificial intelligence; Engineering","score_opus":0.04776571589180073,"score_gpt":0.28914278639383373,"score_spread":0.241377070502033,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1500946169","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016475031,0.009266861,0.031466298,0.06218889,0.03570807,0.00025805802,0.00209559,0.005717194,0.836824],"genre_scores_gemma":[0.16126189,0.0053048246,0.018047115,0.03401149,0.0038834566,0.00030072915,0.0015338963,0.0035056628,0.7721509],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975715,0.0006528144,0.0001514696,0.000593659,0.0006683212,0.00036220253],"domain_scores_gemma":[0.9950858,0.0007741044,0.0005951285,0.0015081434,0.0012143984,0.0008224636],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002170216,0.0008701937,0.0007490934,0.0022795561,0.0045373375,0.0054785754,0.0012465769,0.0027399228,0.12598155],"category_scores_gemma":[0.015776489,0.00038675868,0.0009111666,0.0014327706,0.0030733705,0.006173165,0.006729552,0.003609338,0.046554644],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017834989,0.00005907641,0.002828857,0.00038690216,0.000033026165,0.0011005873,0.00822458,0.00014637121,0.0010978073,0.2280999,0.5556972,0.20214747],"study_design_scores_gemma":[0.000007969404,0.000021670028,0.00035655726,0.000118378055,0.000008313299,0.000409418,0.0006758971,0.000054006054,0.00012666822,0.009073375,0.9891349,0.000012806343],"about_ca_topic_score_codex":0.0047148997,"about_ca_topic_score_gemma":0.005786122,"teacher_disagreement_score":0.12598155,"about_ca_system_score_codex":0.0022720722,"about_ca_system_score_gemma":0.001993403,"threshold_uncertainty_score":0.42145026},"labels":[],"label_agreement":null},{"id":"W1501358284","doi":"","title":"Évaluation et expérimentation de logiciels libres pour la petite et moyenne entreprise","year":2005,"lang":"fr","type":"preprint","venue":"RePEc: Research Papers in Economics","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science; Humanities; Art","score_opus":0.05376283213110975,"score_gpt":0.3424749890882285,"score_spread":0.2887121569571187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1501358284","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76157767,0.013331081,0.15708622,0.0024893153,0.0006194762,0.0026878247,0.003719554,0.012406514,0.046082344],"genre_scores_gemma":[0.65651304,0.005814114,0.29218706,0.000530564,0.00012017334,0.001422727,0.008817315,0.0012113004,0.033383645],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.987528,0.004651875,0.0010849906,0.0011710738,0.0052392897,0.0003248114],"domain_scores_gemma":[0.93758357,0.030920757,0.002808649,0.0050185127,0.02214332,0.001525304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014301838,0.0012507698,0.0008755962,0.004857583,0.0012020676,0.005619805,0.0021576355,0.0014933967,0.0089981165],"category_scores_gemma":[0.056624908,0.0005954936,0.0009524013,0.0043715974,0.0012227107,0.007057213,0.0020693482,0.0016158712,0.0034856712],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024842152,0.0027523923,0.019666582,0.006142229,0.0004513191,0.0006971796,0.006214596,0.012378097,0.03334393,0.012473531,0.01507413,0.8883219],"study_design_scores_gemma":[0.0016285408,0.014158337,0.10229087,0.0049603945,0.0022898174,0.002588901,0.023253364,0.25317702,0.19969936,0.013762023,0.38137457,0.00081681943],"about_ca_topic_score_codex":0.010419002,"about_ca_topic_score_gemma":0.009967431,"teacher_disagreement_score":0.014301838,"about_ca_system_score_codex":0.0026396133,"about_ca_system_score_gemma":0.003596631,"threshold_uncertainty_score":0.07563627},"labels":[],"label_agreement":null},{"id":"W1502857873","doi":"10.1007/978-3-540-27777-4_4","title":"The Role of Process Measurement in Test-Driven Development","year":2004,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Process (computing); Test (biology); Computer science; Geology; Paleontology; Programming language","score_opus":0.01934607856686939,"score_gpt":0.24640339706423026,"score_spread":0.22705731849736085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1502857873","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009755517,0.00508679,0.9679642,0.0012196593,0.00017455567,0.00008436872,0.000054043838,0.0020129909,0.013647955],"genre_scores_gemma":[0.5095053,0.003874906,0.48044285,0.000471563,0.00030106725,0.00035022805,0.00018796032,0.00086033216,0.0040057916],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97874063,0.0107350135,0.0011690157,0.0013776001,0.0074683735,0.00050937117],"domain_scores_gemma":[0.8742857,0.10465837,0.0040315767,0.0111423265,0.005064061,0.000817926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015520367,0.0015809361,0.0013982767,0.00289746,0.0006568622,0.0058227004,0.0029427148,0.0022906906,0.002522475],"category_scores_gemma":[0.07436179,0.0013166324,0.00078917533,0.0032825323,0.00475499,0.011604798,0.0025521624,0.003965418,0.0010291163],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026714097,0.00022163562,0.0039630807,0.00087424956,0.00009437551,0.00011505181,0.0010385753,0.027562043,0.004808386,0.23186775,0.0023264447,0.72686124],"study_design_scores_gemma":[0.00008154062,0.00065721275,0.003092794,0.00088168576,0.00016300404,0.00045221462,0.00031965922,0.34727046,0.02173675,0.60059994,0.02457146,0.00017339247],"about_ca_topic_score_codex":0.0017392584,"about_ca_topic_score_gemma":0.0011926601,"teacher_disagreement_score":0.015520367,"about_ca_system_score_codex":0.0016650979,"about_ca_system_score_gemma":0.0018837568,"threshold_uncertainty_score":0.08208048},"labels":[],"label_agreement":null},{"id":"W1504156788","doi":"10.1109/icse.2015.203","title":"A Unified Framework for the Comprehension of Software's Time","year":2015,"lang":"en","type":"article","venue":"2015 IEEE/ACM 37th IEEE International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software evolution; Program comprehension; Dimension (graph theory); Comprehension; Software engineering; Context (archaeology); Software; Software development; Software construction; Software system; Programming language; Mathematics","score_opus":0.10157726757901188,"score_gpt":0.3405022812091452,"score_spread":0.23892501363013335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1504156788","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030075994,0.0007194624,0.97771204,0.0036273298,0.00016050805,0.00008873823,0.00016665479,0.0006453828,0.013872281],"genre_scores_gemma":[0.19472638,0.0012456689,0.79517454,0.0010158962,0.00065241865,0.0008031811,0.00055354764,0.000893606,0.004934754],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9923983,0.0035743508,0.0008047304,0.0013766282,0.001259724,0.0005863634],"domain_scores_gemma":[0.9886996,0.0058139884,0.0009125406,0.0020980341,0.0018274864,0.00064842444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012451001,0.0021137372,0.0019727496,0.0062990854,0.0033789047,0.012125903,0.004896992,0.0057062884,0.008398519],"category_scores_gemma":[0.018131128,0.0015149212,0.0049596536,0.004157207,0.017422514,0.032660596,0.006737662,0.00885599,0.0020240308],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000052439477,0.0000071041272,0.0000557364,0.00003275193,0.0000069404205,0.000047360645,0.0012159132,0.00070499285,0.00020783133,0.9944602,0.00044571303,0.0028102112],"study_design_scores_gemma":[0.000012885184,0.000019725006,0.00008852882,0.00007278551,0.000014184062,0.00009307757,0.00049144606,0.0059086625,0.00027350493,0.97747016,0.0155313965,0.000023540966],"about_ca_topic_score_codex":0.006162268,"about_ca_topic_score_gemma":0.0036075003,"teacher_disagreement_score":0.012451001,"about_ca_system_score_codex":0.004829788,"about_ca_system_score_gemma":0.0051918477,"threshold_uncertainty_score":0.06584799},"labels":[],"label_agreement":null},{"id":"W1504662635","doi":"10.1007/978-3-642-36089-3_12","title":"Grammatical Inference in Software Engineering: An Overview of the State of the Art","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Grammar induction; Computer science; Inference; Rule-based machine translation; Grammar; Artificial intelligence; Natural language processing; Software; Variety (cybernetics); Programming language; Class (philosophy); Process (computing); Linguistics","score_opus":0.03433466987507978,"score_gpt":0.2773666441244039,"score_spread":0.24303197424932416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1504662635","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026095547,0.5976459,0.30725917,0.009442481,0.0017035112,0.000090431866,0.00042934576,0.0016016443,0.07921801],"genre_scores_gemma":[0.06251202,0.5700957,0.3207753,0.0038896743,0.006181714,0.00020622807,0.001867667,0.0018681871,0.032603476],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982406,0.00070476363,0.00015120102,0.00027746943,0.00054441945,0.00008149623],"domain_scores_gemma":[0.99542326,0.0037643209,0.000096205265,0.00030176013,0.00035669032,0.000057685764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032205626,0.0010028301,0.0016192464,0.0032086512,0.00066681515,0.004760289,0.0023298718,0.002516768,0.009545666],"category_scores_gemma":[0.0055102566,0.00077802263,0.0010695746,0.004797475,0.0033902314,0.0075670574,0.0019405892,0.00349973,0.0052330745],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000314671,0.00007610338,0.00026929684,0.0030378106,0.000057746965,0.000100997684,0.00069956284,0.0023074436,0.0015529844,0.27169088,0.036625575,0.6835502],"study_design_scores_gemma":[0.000011463326,0.000029707953,0.0005569527,0.0016937562,0.00004483043,0.00044801048,0.00021219699,0.006911257,0.0015594019,0.4902358,0.49825162,0.000045027013],"about_ca_topic_score_codex":0.0016015384,"about_ca_topic_score_gemma":0.0023270028,"teacher_disagreement_score":0.009545666,"about_ca_system_score_codex":0.0014531479,"about_ca_system_score_gemma":0.0023089016,"threshold_uncertainty_score":0.031933367},"labels":[],"label_agreement":null},{"id":"W1505022629","doi":"10.1109/metrics.2005.3","title":"A Model for Performance Management and Estimation","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Code smell; Code refactoring; Computer science; Software engineering; Interpretation (philosophy); Software; Java; Set (abstract data type); Intuition; Source code; Static program analysis; Software system; Software development; Software quality; Programming language; Artificial intelligence","score_opus":0.02590379943914148,"score_gpt":0.27115083727410794,"score_spread":0.24524703783496646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1505022629","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0056005553,0.0007784434,0.95791334,0.0027724071,0.00024298485,0.00024515745,0.001715733,0.0015474943,0.029183896],"genre_scores_gemma":[0.5314841,0.0038405873,0.33837396,0.00088723004,0.0008907969,0.0024330595,0.0042623486,0.0009820919,0.116845824],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99512416,0.0016092653,0.00027283447,0.0013089649,0.001147661,0.0005371725],"domain_scores_gemma":[0.9915005,0.004758938,0.00079941715,0.001081072,0.0016100662,0.0002500207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006162482,0.002396481,0.001660925,0.0026377924,0.0011071697,0.0065582413,0.005663344,0.005156704,0.030206427],"category_scores_gemma":[0.022242768,0.0012193838,0.002084699,0.003669445,0.0017902964,0.010049893,0.0023372953,0.004069092,0.013548684],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013252885,0.00013618074,0.0020164934,0.00018202561,0.00012739962,0.00018988806,0.00030685586,0.46389475,0.00050048,0.46956563,0.011291275,0.051656526],"study_design_scores_gemma":[0.000039929047,0.00005663897,0.00064199156,0.000059091555,0.00004719779,0.000090696056,0.00007841936,0.78838855,0.00020044172,0.18816558,0.022180008,0.00005146503],"about_ca_topic_score_codex":0.01816307,"about_ca_topic_score_gemma":0.00875762,"teacher_disagreement_score":0.030206427,"about_ca_system_score_codex":0.00462104,"about_ca_system_score_gemma":0.0029030198,"threshold_uncertainty_score":0.101050496},"labels":[],"label_agreement":null},{"id":"W1507777432","doi":"10.1609/aimag.v32i2.2348","title":"AI‐Based Software Defect Predictors: Applications and Benefits in a Case Study","year":2011,"lang":"en","type":"article","venue":"AI Magazine","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Process (computing); Software; Computer science; Software bug; Code (set theory); Reliability engineering; Software engineering; Software inspection; Predictive modelling; Software development; Machine learning; Artificial intelligence; Software quality; Engineering; Operating system; Programming language","score_opus":0.030335932235699782,"score_gpt":0.272739207723545,"score_spread":0.2424032754878452,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1507777432","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9515102,0.0002682753,0.043244682,0.00082538684,0.000023247987,0.00020404677,0.00028878934,0.00052310055,0.0031124207],"genre_scores_gemma":[0.9482369,0.00018632782,0.05017878,0.00004045505,0.000010737916,0.000094576215,0.00018596405,0.000037811456,0.0010285524],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99852496,0.0009562677,0.00007975546,0.00012207458,0.00025266423,0.00006427843],"domain_scores_gemma":[0.9807334,0.016116414,0.00056793785,0.0008962076,0.0013607811,0.00032521947],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027065123,0.0007458736,0.00041140057,0.0014267518,0.00047111575,0.0007249187,0.0011494698,0.0014165929,0.0014099556],"category_scores_gemma":[0.010910121,0.00031319904,0.0004756611,0.0015694802,0.0006966528,0.0009460147,0.0006939909,0.00085530063,0.00021899956],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00152244,0.005505382,0.19524619,0.00053177366,0.00021593265,0.0066088033,0.0027621922,0.48828793,0.007497637,0.00723068,0.0046147048,0.2799763],"study_design_scores_gemma":[0.00013772088,0.0010219576,0.016333804,0.000038661132,0.00006465006,0.0007295179,0.000779648,0.9713164,0.0053868666,0.0023356362,0.0018092386,0.00004580845],"about_ca_topic_score_codex":0.013246236,"about_ca_topic_score_gemma":0.015245055,"teacher_disagreement_score":0.013246236,"about_ca_system_score_codex":0.0010117857,"about_ca_system_score_gemma":0.0006585323,"threshold_uncertainty_score":0.02633828},"labels":[],"label_agreement":null},{"id":"W1508206648","doi":"10.1109/ipcc.1994.347497","title":"Shrink-wrapping the formula for better information: end user vs. MIS computer documentation preferences","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Documentation; Software documentation; Product (mathematics); Software; IBM; Computer science; World Wide Web; End user; Software engineering; Software development; Operating system; Software development process","score_opus":0.025060131367528465,"score_gpt":0.25272415067083936,"score_spread":0.2276640193033109,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1508206648","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98691934,0.00035065852,0.00067482033,0.0017264314,0.000020327456,0.0000117739455,0.00005680278,0.000012098093,0.0102277],"genre_scores_gemma":[0.99739516,0.00020538321,0.00037364272,0.00056110555,0.000017317818,0.000009652891,0.000042232517,0.000016025346,0.0013793379],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9959352,0.0024350015,0.00033219083,0.00016259162,0.00078899047,0.00034597013],"domain_scores_gemma":[0.97076744,0.01882767,0.0030752323,0.00090753235,0.0036159353,0.0028062009],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005466918,0.00012166213,0.0002174561,0.0012815136,0.0008072781,0.0028438831,0.00025087708,0.0009716474,0.0068311975],"category_scores_gemma":[0.036905237,0.00015135047,0.00026468388,0.0008395271,0.0009632802,0.0033175098,0.00097406656,0.0008087788,0.000984294],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020527039,0.00058244367,0.7518296,0.00026198468,0.000092423004,0.00082245964,0.077573344,0.00021845031,0.0022014214,0.006167771,0.011786462,0.14641088],"study_design_scores_gemma":[0.00008325942,0.0013688469,0.74161243,0.00036355952,0.000100683676,0.002974574,0.22362874,0.0021403986,0.0021639846,0.0039396053,0.021505421,0.00011858075],"about_ca_topic_score_codex":0.0016053037,"about_ca_topic_score_gemma":0.0022661972,"teacher_disagreement_score":0.0068311975,"about_ca_system_score_codex":0.00047070504,"about_ca_system_score_gemma":0.0003400688,"threshold_uncertainty_score":0.028912187},"labels":[],"label_agreement":null},{"id":"W1508407965","doi":"10.1007/978-1-4615-0429-0_12","title":"Condensing Uncertainty via Incremental Treatment Learning","year":2003,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Incremental learning; Computer science; Artificial intelligence","score_opus":0.02693826942481265,"score_gpt":0.2567579758442118,"score_spread":0.22981970641939914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1508407965","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024517167,0.0006780672,0.98823804,0.00033793284,0.00010459814,0.000026581729,0.00005491755,0.0005082648,0.007599973],"genre_scores_gemma":[0.31016588,0.0023370576,0.6451815,0.0006585696,0.0006815393,0.00034685983,0.00066251995,0.0007798499,0.039186165],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99916625,0.00023485186,0.000035779172,0.00013061342,0.00037567536,0.00005694548],"domain_scores_gemma":[0.9976463,0.0013796326,0.00008573746,0.00054361665,0.00029263448,0.000051983112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013513383,0.0007170509,0.0010921645,0.0008246711,0.00053845654,0.0012273454,0.0018282458,0.00091617357,0.008727467],"category_scores_gemma":[0.007112153,0.0005114721,0.0008723214,0.0012985374,0.0015035679,0.0037188593,0.0024478203,0.0037111526,0.0015203346],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007248675,0.00008276668,0.00023597041,0.0002238952,0.000068441885,0.00005883905,0.00014467262,0.15231706,0.002593125,0.27682558,0.0136619955,0.5537152],"study_design_scores_gemma":[0.000016854425,0.000036547033,0.00012053363,0.000039490133,0.00003223524,0.00005683098,0.000024798534,0.5041323,0.0034624715,0.47912455,0.012930415,0.000023025932],"about_ca_topic_score_codex":0.0010665043,"about_ca_topic_score_gemma":0.0018246709,"teacher_disagreement_score":0.008727467,"about_ca_system_score_codex":0.00095833885,"about_ca_system_score_gemma":0.000952892,"threshold_uncertainty_score":0.029196262},"labels":[],"label_agreement":null},{"id":"W1508590353","doi":"10.1002/smr.1597","title":"An empirical study of faults in late propagation clone genealogies","year":2013,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"clone (Java method); Commit; Software evolution; Biology; Software; Computer science; Software system; Genetics; Programming language; Gene; Database","score_opus":0.021822808823423523,"score_gpt":0.3192357228767848,"score_spread":0.2974129140533613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1508590353","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99901795,0.000050831284,0.0005625872,0.000031840304,0.0000011304073,0.000009089972,0.000088290566,0.000013376133,0.00022491795],"genre_scores_gemma":[0.999433,0.00001684911,0.00031151142,0.000006776364,0.0000016824775,0.0000075587586,0.00013217321,0.0000050967783,0.00008531221],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99660146,0.0013153349,0.00037229116,0.00063086534,0.0008571014,0.0002229563],"domain_scores_gemma":[0.80820876,0.13347666,0.035997663,0.009389602,0.010561857,0.0023655035],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.003855944,0.00022647863,0.00025942648,0.0029630612,0.000564685,0.0010760177,0.0010059439,0.000842139,0.0010562501],"category_scores_gemma":[0.075328745,0.0002921294,0.00025722926,0.002551991,0.001523779,0.0022833082,0.0009027416,0.0011594794,0.00020444198],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009925189,0.000096638214,0.98862326,0.000028086362,0.000042011146,0.00024120523,0.001611973,0.0013260139,0.00047336592,0.00033909234,0.00015173024,0.006967411],"study_design_scores_gemma":[0.000013561523,0.00028501693,0.9740397,0.000025404152,0.000029845327,0.0010900961,0.0030222805,0.018702516,0.0008188405,0.0012700157,0.0006790375,0.000023749271],"about_ca_topic_score_codex":0.0039454442,"about_ca_topic_score_gemma":0.004635538,"teacher_disagreement_score":0.99614406,"about_ca_system_score_codex":0.00077494886,"about_ca_system_score_gemma":0.00041650384,"threshold_uncertainty_score":0.020392418},"labels":[],"label_agreement":null},{"id":"W1508879959","doi":"10.1109/wpc.2005.38","title":"Theories, methods and tools in program comprehension: past, present and future","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":207,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Program comprehension; Comprehension; Computer science; Construct (python library); Context (archaeology); Cognition; Key (lock); Data science; Cognitive science; Software; Software system; Psychology; Programming language","score_opus":0.02900956144356537,"score_gpt":0.35894913008395185,"score_spread":0.3299395686403865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1508879959","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015932955,0.50685734,0.2768225,0.14328885,0.0028668095,0.0002573508,0.00012202934,0.0011997687,0.052652303],"genre_scores_gemma":[0.13668959,0.39465982,0.43736044,0.009435717,0.003742792,0.0011536638,0.0002598331,0.00043409647,0.016264103],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98612416,0.009022572,0.0009997555,0.00086154067,0.0025393378,0.00045267318],"domain_scores_gemma":[0.9088536,0.077731974,0.0032934467,0.0041532163,0.004404352,0.0015634117],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.038081996,0.0016712261,0.001360218,0.012031087,0.0019824782,0.016262632,0.0032786685,0.005498847,0.00540428],"category_scores_gemma":[0.038268264,0.0010597621,0.0013954864,0.010830028,0.020037003,0.042202927,0.003719593,0.0060512037,0.0015061845],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000616336,0.00035358552,0.0041560032,0.0028034435,0.000056707504,0.00015120977,0.008914736,0.0013554253,0.00050618773,0.4383923,0.011864955,0.5313839],"study_design_scores_gemma":[0.000042435084,0.00016166289,0.0027065205,0.0047802683,0.000047607686,0.00043720202,0.012128939,0.004990735,0.0012617448,0.78051716,0.19278786,0.00013792937],"about_ca_topic_score_codex":0.0030369398,"about_ca_topic_score_gemma":0.002185124,"teacher_disagreement_score":0.038081996,"about_ca_system_score_codex":0.00457129,"about_ca_system_score_gemma":0.007266076,"threshold_uncertainty_score":0.2013992},"labels":[],"label_agreement":null},{"id":"W1508909554","doi":"10.1007/978-3-319-00948-3_3","title":"Exploring a Model-Oriented and Executable Syntax for UML Attributes","year":2013,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Unified Modeling Language; Class diagram; Programming language; Executable; Object Constraint Language; Java; Code generation; Abstract syntax; Syntax; Object-oriented programming; Source code; Modeling language; Applications of UML; Semantics (computer science); Artificial intelligence; Software","score_opus":0.3161396005265878,"score_gpt":0.3637513490700713,"score_spread":0.04761174854348349,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1508909554","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005982178,0.00023341134,0.98434055,0.00071653444,0.000045263485,0.00006642935,0.00026633643,0.0010026024,0.007346653],"genre_scores_gemma":[0.06461674,0.0006014301,0.9278485,0.00018314179,0.000038457783,0.00011650496,0.0009227866,0.0012294087,0.0044429074],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9988887,0.00044772457,0.00013113272,0.0001746564,0.00029513417,0.00006262591],"domain_scores_gemma":[0.99784684,0.0012684413,0.00014492507,0.00030998097,0.0003589236,0.000070976625],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023204144,0.00087314175,0.0004844197,0.001088877,0.0010141216,0.005931668,0.0017540458,0.0013227661,0.007006014],"category_scores_gemma":[0.006019937,0.0009814416,0.0012830757,0.0018049482,0.0021198338,0.008641726,0.002553521,0.0032083383,0.0018928753],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022041233,0.000021936425,0.00024911299,0.0001627675,0.0000069915986,0.000106759166,0.0020225884,0.0027795825,0.0037018675,0.9398017,0.0026823208,0.048442326],"study_design_scores_gemma":[0.000026818796,0.000044499004,0.00019955797,0.000393539,0.00004028809,0.00043144068,0.0010037958,0.055204455,0.0106392875,0.77825713,0.15370256,0.000056586632],"about_ca_topic_score_codex":0.0024230094,"about_ca_topic_score_gemma":0.0034374243,"teacher_disagreement_score":0.007006014,"about_ca_system_score_codex":0.001363115,"about_ca_system_score_gemma":0.0021874937,"threshold_uncertainty_score":0.02343744},"labels":[],"label_agreement":null},{"id":"W1509092525","doi":"10.1023/a:1022697422240","title":"The Migration of Multi-tier E-commerce Applications to an Enterprise Java Environment","year":2003,"lang":"en","type":"article","venue":"Information Systems Frontiers","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Toronto; University of Windsor; IBM (Canada)","funders":"Japan Society for the Promotion of Science","keywords":"IBM; Computer science; Legacy system; JavaBeans; Software engineering; Suite; Legacy code; Database; Java; Programming language; Software","score_opus":0.012868648569035461,"score_gpt":0.24185270636003212,"score_spread":0.22898405779099665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1509092525","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9018245,0.0001775717,0.082092956,0.00059656193,0.00025784888,0.00017752394,0.00013228794,0.0070532574,0.0076876017],"genre_scores_gemma":[0.9258561,0.00013234049,0.062331382,0.0003409077,0.000026340882,0.000044202203,0.00049897167,0.00083129795,0.009938485],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99910223,0.00016655552,0.000098094715,0.00014516234,0.0002500421,0.00023793674],"domain_scores_gemma":[0.9972505,0.00048017933,0.00017535602,0.0010939003,0.0005715152,0.0004285886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014205958,0.0003368595,0.00026732805,0.00035162712,0.00083343836,0.0018717679,0.0013532187,0.0008036497,0.0018581271],"category_scores_gemma":[0.0050639873,0.00052569003,0.00040253397,0.00044256492,0.00044721673,0.0017911615,0.001745202,0.0013167133,0.0006134242],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006488677,0.0030316808,0.0834501,0.00026299298,0.00018133083,0.0050999667,0.006790309,0.027128898,0.32517746,0.030573642,0.019587485,0.4922275],"study_design_scores_gemma":[0.0005524183,0.0014454813,0.07168664,0.00011604223,0.00033937115,0.0030186789,0.0033868589,0.5928204,0.24003878,0.016157126,0.07020913,0.000229048],"about_ca_topic_score_codex":0.0070566754,"about_ca_topic_score_gemma":0.0067503937,"teacher_disagreement_score":0.0070566754,"about_ca_system_score_codex":0.00048669125,"about_ca_system_score_gemma":0.0009662165,"threshold_uncertainty_score":0.014031231},"labels":[],"label_agreement":null},{"id":"W1509390953","doi":"","title":"Proceedings of the 4th International Workshop on Software Clones","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Code refactoring; Software engineering; Computer science; clone (Java method); Software; Restructuring; Software construction; Software development; Software system; Software quality; Programming language; Political science","score_opus":0.014210566112329953,"score_gpt":0.2624778953677662,"score_spread":0.24826732925543626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1509390953","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041081958,0.07283739,0.4717942,0.07323867,0.06912082,0.0017754928,0.0017770157,0.0047148154,0.26365963],"genre_scores_gemma":[0.14637855,0.045205064,0.25654125,0.01116055,0.016765157,0.0020320858,0.009303101,0.0037837762,0.50883055],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99420446,0.0020442796,0.00045107526,0.0009286417,0.0017431822,0.00062842655],"domain_scores_gemma":[0.9914226,0.002755567,0.0003014582,0.0015111361,0.0027247132,0.0012845624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008094373,0.0015434795,0.0014538959,0.0024988283,0.0023242468,0.0073774094,0.0028651732,0.004071481,0.04391019],"category_scores_gemma":[0.013688177,0.001021168,0.0020767713,0.0021042689,0.002033697,0.009363893,0.0066969097,0.0044065616,0.01187894],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000542809,0.00036290393,0.0012969667,0.00075931486,0.0000881168,0.0008031421,0.0031727736,0.0024621857,0.0066508665,0.052378066,0.47962502,0.45185786],"study_design_scores_gemma":[0.000068076784,0.00014639004,0.00093129557,0.0004959438,0.000059390986,0.000676307,0.0008481795,0.002609814,0.0021585366,0.018969761,0.9729934,0.000042808333],"about_ca_topic_score_codex":0.0021577075,"about_ca_topic_score_gemma":0.003356789,"teacher_disagreement_score":0.04391019,"about_ca_system_score_codex":0.0021472566,"about_ca_system_score_gemma":0.0030487217,"threshold_uncertainty_score":0.14689416},"labels":[],"label_agreement":null},{"id":"W1509675159","doi":"10.1023/a:1025801405075","title":"Agile Parsing in TXL","year":2003,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; University of Waterloo; Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Parsing; Grammar; Programming language; Agile software development; Natural language processing; Task (project management); Software engineering; Artificial intelligence; Set (abstract data type); Top-down parsing; Linguistics; Engineering; Systems engineering","score_opus":0.009901452217189218,"score_gpt":0.23796285214521737,"score_spread":0.22806139992802815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1509675159","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0253006,0.00026886287,0.84656197,0.00095979404,0.00025026646,0.000120895216,0.0023497718,0.09605532,0.02813252],"genre_scores_gemma":[0.33361784,0.00038345705,0.58296734,0.0010416482,0.00017058152,0.00022049526,0.010135211,0.033411484,0.038051926],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99702626,0.0008475254,0.0002945092,0.00058156403,0.00091250753,0.00033766805],"domain_scores_gemma":[0.99443746,0.002121097,0.00022128403,0.0023824696,0.00071473815,0.00012289081],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027000536,0.00085708936,0.00065181684,0.0015304317,0.0008978764,0.0038788714,0.0017208124,0.0012805125,0.019647926],"category_scores_gemma":[0.0075618923,0.0012219296,0.0009379484,0.0020643657,0.001162561,0.005828088,0.003076443,0.0024534736,0.009419461],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006267497,0.00020422965,0.0054505696,0.0005055504,0.00009360737,0.00091949507,0.0018481865,0.017560437,0.016879318,0.22872612,0.10410167,0.62308407],"study_design_scores_gemma":[0.00017455591,0.00014541208,0.0017628093,0.00029727782,0.00015595902,0.00087325874,0.00075555686,0.24738191,0.11716056,0.29073218,0.34038007,0.00018045865],"about_ca_topic_score_codex":0.0031058767,"about_ca_topic_score_gemma":0.0036381085,"teacher_disagreement_score":0.019647926,"about_ca_system_score_codex":0.00095087907,"about_ca_system_score_gemma":0.0019401798,"threshold_uncertainty_score":0.06572884},"labels":[],"label_agreement":null},{"id":"W1510033441","doi":"10.1007/978-3-540-68073-4_35","title":"Lightweight, Semi-automated Enactment of Pragmatic-Reuse Plans","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Computer science; Task (project management); Plan (archaeology); Software engineering; Process (computing); Code reuse; Code (set theory); Focus (optics); Source code; Human–computer interaction; Programming language; Software; Systems engineering","score_opus":0.013655363210843057,"score_gpt":0.24980066308453872,"score_spread":0.23614529987369565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1510033441","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014889665,0.00010333056,0.9685981,0.00019909718,0.00004760675,0.0002941952,0.00019665307,0.012241418,0.003429854],"genre_scores_gemma":[0.25566447,0.00014278048,0.73591936,0.00013109123,0.00004128432,0.00027966007,0.0009857219,0.0019748188,0.0048608226],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98985296,0.0032999704,0.00067084376,0.0009298953,0.004498342,0.00074789906],"domain_scores_gemma":[0.9757858,0.011412622,0.0011803454,0.008742859,0.002426023,0.00045237254],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006205031,0.0013850755,0.0012969286,0.0016052687,0.0012297865,0.0039433553,0.003061457,0.0017265109,0.006960156],"category_scores_gemma":[0.030460551,0.0015389161,0.0020540925,0.0010208698,0.0023892578,0.004427207,0.0074989283,0.0027229346,0.002925085],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010162286,0.0005213613,0.003970472,0.0009094255,0.00026878336,0.0008311846,0.002977532,0.07903781,0.045358226,0.14809617,0.012786795,0.70422596],"study_design_scores_gemma":[0.000153588,0.0001964068,0.0013514254,0.00019118545,0.00016518398,0.00046750894,0.00075158774,0.6921553,0.074911945,0.2030247,0.026485037,0.00014610129],"about_ca_topic_score_codex":0.0054891314,"about_ca_topic_score_gemma":0.008169525,"teacher_disagreement_score":0.006960156,"about_ca_system_score_codex":0.0011480477,"about_ca_system_score_gemma":0.004321527,"threshold_uncertainty_score":0.032815754},"labels":[],"label_agreement":null},{"id":"W1512221544","doi":"","title":"Robust multilingual parsing using island grammars","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Programming language; Parsing; Parsing expression grammar; S-attributed grammar; Rule-based machine translation; L-attributed grammar; Natural language processing; Artificial intelligence; Bottom-up parsing; Syntax; Top-down parsing; Top-down parsing language; Context-free grammar","score_opus":0.06253956424611067,"score_gpt":0.29037058015684414,"score_spread":0.22783101591073346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1512221544","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007772125,0.0001336937,0.97137636,0.00020181449,0.00005978548,0.00007366788,0.0006045575,0.013240836,0.006537156],"genre_scores_gemma":[0.20035928,0.0003347185,0.77593845,0.0002031465,0.000099594065,0.00022289586,0.003963313,0.011472193,0.0074064564],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969146,0.0009038548,0.0003335224,0.00085501996,0.00076902757,0.00022404235],"domain_scores_gemma":[0.9925339,0.0026849178,0.0003189448,0.0029833058,0.0013069032,0.00017188184],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027442768,0.0011064063,0.0011179583,0.0019652206,0.001352427,0.003300052,0.002253217,0.001304106,0.008162184],"category_scores_gemma":[0.009671814,0.0011818092,0.0015308361,0.00186093,0.0017839983,0.004966954,0.005983072,0.0024032746,0.005441442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037558886,0.00022792742,0.0044560004,0.0007493157,0.000281324,0.0029497507,0.0037886007,0.07434065,0.04854918,0.28449807,0.036298294,0.5434853],"study_design_scores_gemma":[0.000073325435,0.00006802734,0.001292772,0.00016039332,0.00015376275,0.0012238174,0.00064481265,0.35761094,0.054014195,0.4839445,0.10060712,0.00020636256],"about_ca_topic_score_codex":0.0033703218,"about_ca_topic_score_gemma":0.0049292347,"teacher_disagreement_score":0.008162184,"about_ca_system_score_codex":0.0007520362,"about_ca_system_score_gemma":0.0020504654,"threshold_uncertainty_score":0.027305186},"labels":[],"label_agreement":null},{"id":"W1512615396","doi":"10.1002/smr.1662","title":"Big data clone detection using classical detectors: an exploratory study","year":2014,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Saskatchewan","funders":"","keywords":"Computer science; Scalability; Big data; Source code; Code (set theory); clone (Java method); Data mining; Data science; Database; Programming language","score_opus":0.07457598398553884,"score_gpt":0.3182856481568901,"score_spread":0.24370966417135126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1512615396","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9474265,0.0007315926,0.0475673,0.0003483577,0.000032014494,0.00058379374,0.0009030233,0.00075470714,0.0016528153],"genre_scores_gemma":[0.9297568,0.00024788402,0.067676894,0.00012163643,0.000034176606,0.0002762415,0.0013613595,0.00012854904,0.00039662066],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.985274,0.0074578132,0.00074671896,0.001844381,0.004210997,0.00046610532],"domain_scores_gemma":[0.8066698,0.15384364,0.0089576775,0.014876306,0.014196757,0.0014558705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014691922,0.001065088,0.00073353125,0.005555054,0.00096253527,0.0021071224,0.002528153,0.0012946911,0.0004681821],"category_scores_gemma":[0.06631155,0.00055089325,0.0012875777,0.0033060736,0.0017705323,0.0031946416,0.002114549,0.001918138,0.00030901257],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019109286,0.005127853,0.63149756,0.0023520654,0.0011065367,0.0022653765,0.0102344025,0.038335495,0.021246148,0.008557207,0.009568627,0.26779783],"study_design_scores_gemma":[0.00032658607,0.004704618,0.2176282,0.00060539605,0.00063147425,0.00439697,0.011980883,0.6684609,0.05301902,0.018851181,0.019059755,0.00033508445],"about_ca_topic_score_codex":0.0023577553,"about_ca_topic_score_gemma":0.0026682573,"teacher_disagreement_score":0.014691922,"about_ca_system_score_codex":0.0010199001,"about_ca_system_score_gemma":0.00086472236,"threshold_uncertainty_score":0.077699244},"labels":[],"label_agreement":null},{"id":"W1512857087","doi":"10.1109/iwpse.2004.3","title":"Aiding comprehension of cloning through categorization","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":68,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"clone (Java method); Program comprehension; Computer science; False positive paradox; Software maintenance; Cloning (programming); Source code; Function (biology); Software; Programming language; Categorization; Filter (signal processing); Software system; Artificial intelligence; Biology; Genetics; Gene","score_opus":0.02673080414723087,"score_gpt":0.2748554369435063,"score_spread":0.24812463279627542,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1512857087","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2644152,0.00081959256,0.7136276,0.0019674208,0.000093862356,0.0004469667,0.000511428,0.011842135,0.0062758145],"genre_scores_gemma":[0.49984342,0.00045448382,0.4944472,0.00049411727,0.000061125786,0.000256577,0.0014115089,0.00053678185,0.0024948718],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9894348,0.0048372475,0.0010372641,0.0016896668,0.0025024507,0.00049865944],"domain_scores_gemma":[0.90224105,0.06643207,0.0064358483,0.010790504,0.013074989,0.0010256082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010640964,0.0011793195,0.0013710774,0.0071467706,0.0012996502,0.0053151944,0.002950612,0.002860216,0.0026558596],"category_scores_gemma":[0.07976091,0.0006707451,0.0008641029,0.0035292895,0.0014429498,0.013186312,0.0038479392,0.0018778542,0.0012363625],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000537604,0.00048099735,0.047875743,0.00108631,0.000100116464,0.0007304499,0.02604071,0.006561411,0.05946317,0.024224155,0.007536348,0.825363],"study_design_scores_gemma":[0.000170024,0.0014396197,0.056172103,0.0010900762,0.00041823008,0.0056808763,0.020320531,0.4790393,0.14036858,0.18541038,0.109285094,0.00060514547],"about_ca_topic_score_codex":0.0024513109,"about_ca_topic_score_gemma":0.002318404,"teacher_disagreement_score":0.010640964,"about_ca_system_score_codex":0.0014202826,"about_ca_system_score_gemma":0.0022583823,"threshold_uncertainty_score":0.056275487},"labels":[],"label_agreement":null},{"id":"W1513330488","doi":"","title":"Applying data mining to software maintenance records","year":2003,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Software quality; Data mining; Software maintenance; Software; Quality (philosophy); Software system; Software engineering; Software development; Data science; Programming language","score_opus":0.17939863254285643,"score_gpt":0.41542539815001706,"score_spread":0.23602676560716063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1513330488","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5895217,0.0019284429,0.39099473,0.0018863071,0.0001257037,0.00056872383,0.010852909,0.0025677725,0.00155365],"genre_scores_gemma":[0.7006856,0.0007570326,0.2829013,0.00015755925,0.000071056325,0.0002829356,0.014555595,0.00006273816,0.0005262197],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956326,0.0015748898,0.00075910473,0.0009748534,0.00087270164,0.00018578453],"domain_scores_gemma":[0.9632334,0.02764221,0.0020957193,0.00416164,0.002549587,0.00031741214],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005662944,0.0011468672,0.0014520724,0.0080297645,0.001136104,0.0022522428,0.0023837357,0.0015732283,0.00045781163],"category_scores_gemma":[0.035535,0.00074600376,0.0019029885,0.008688647,0.0006694162,0.0023692811,0.0012401219,0.0019594447,0.00039506546],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009530875,0.0014801446,0.19085826,0.0009429288,0.0014271479,0.0009245066,0.0014777518,0.24936935,0.0038080006,0.0032299887,0.0040949564,0.5414339],"study_design_scores_gemma":[0.00007966432,0.00033972348,0.023000136,0.00014969175,0.00032461158,0.00047311155,0.00085125567,0.94772875,0.006012616,0.016217181,0.0047501717,0.000073071744],"about_ca_topic_score_codex":0.01215812,"about_ca_topic_score_gemma":0.014942874,"teacher_disagreement_score":0.01215812,"about_ca_system_score_codex":0.0012050127,"about_ca_system_score_gemma":0.0014811673,"threshold_uncertainty_score":0.02994889},"labels":[],"label_agreement":null},{"id":"W1514984630","doi":"10.1007/978-3-642-01818-3_28","title":"Exploratory Analysis of Co-Change Graphs for Code Refactoring","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Code refactoring; Cluster analysis; Software; Software system; Probabilistic logic; Theoretical computer science; Data mining; Software engineering; Distributed computing; Programming language; Artificial intelligence","score_opus":0.05951093469987863,"score_gpt":0.30765371465937946,"score_spread":0.24814277995950085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1514984630","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7459925,0.0011322572,0.19341117,0.00048798375,0.00010970784,0.00082614983,0.02917561,0.02060753,0.008257036],"genre_scores_gemma":[0.7709837,0.00036628632,0.19844252,0.000048515776,0.000041298877,0.00036113377,0.025118487,0.0018275244,0.0028105609],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99804544,0.00039699685,0.0001415561,0.00040717237,0.0008557719,0.00015304636],"domain_scores_gemma":[0.95981044,0.029454293,0.0022832851,0.0032936234,0.0045954674,0.00056275364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001900201,0.0006346751,0.00040493786,0.009193016,0.000730572,0.001430363,0.0010121583,0.0006081635,0.0034861467],"category_scores_gemma":[0.018471997,0.0002586051,0.0008337569,0.007717362,0.00043313348,0.0012987055,0.0010200245,0.0010355001,0.0008343378],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012244647,0.0009937809,0.112567976,0.0024217665,0.00044324296,0.0017426736,0.011509114,0.015965601,0.056954097,0.010639176,0.01918781,0.7663503],"study_design_scores_gemma":[0.0001590689,0.001155433,0.39722848,0.00078940805,0.00064713554,0.0034569732,0.0083948355,0.40960735,0.0851927,0.033515226,0.059481256,0.00037213328],"about_ca_topic_score_codex":0.004280094,"about_ca_topic_score_gemma":0.008015788,"teacher_disagreement_score":0.009193016,"about_ca_system_score_codex":0.00052581733,"about_ca_system_score_gemma":0.0009712836,"threshold_uncertainty_score":0.011662364},"labels":[],"label_agreement":null},{"id":"W1516350265","doi":"10.1109/iwsc.2015.7069884","title":"Performance impact of lazy deletion in metric trees for incremental clone analysis","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Computer science; Metric (unit); Software maintenance; Process (computing); Software; Programming language; Software system; Engineering; Biology; Operations management; Gene","score_opus":0.03738886913777351,"score_gpt":0.3228812769442517,"score_spread":0.2854924078064782,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1516350265","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8046159,0.002472,0.16053596,0.0005473869,0.00029806068,0.00016380977,0.00063150044,0.02902501,0.0017104144],"genre_scores_gemma":[0.884186,0.0002280566,0.11317661,0.000101303114,0.000043229997,0.00006223742,0.00086433196,0.00065143727,0.00068667566],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99055326,0.0030083114,0.00084623176,0.0018431896,0.0028887491,0.0008602587],"domain_scores_gemma":[0.9476744,0.033891995,0.002423188,0.008976268,0.005034123,0.0019999442],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070230495,0.0011757491,0.0013960865,0.0021425744,0.0011029588,0.0019885558,0.0032644272,0.0011757907,0.0012389796],"category_scores_gemma":[0.038365044,0.0006043369,0.00089059415,0.0026619658,0.0010561584,0.0057960027,0.0015443601,0.0013465553,0.00063035457],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0060819653,0.00085168076,0.055188324,0.00080939726,0.00049128954,0.00056115055,0.0016073469,0.17035194,0.11375069,0.0077106757,0.00924474,0.63335085],"study_design_scores_gemma":[0.00019178625,0.0017596678,0.011664006,0.000044494336,0.00021637432,0.00048023515,0.00029938266,0.91760147,0.05991758,0.0043101856,0.0033540702,0.00016072382],"about_ca_topic_score_codex":0.010201093,"about_ca_topic_score_gemma":0.0072350954,"teacher_disagreement_score":0.010201093,"about_ca_system_score_codex":0.00167811,"about_ca_system_score_gemma":0.0023732516,"threshold_uncertainty_score":0.03714192},"labels":[],"label_agreement":null},{"id":"W1517245390","doi":"10.1007/978-3-642-20677-1_17","title":"A Test-Driven Approach for Extracting Libraries of Reusable Components from Existing Applications","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Code refactoring; Computer science; Agile software development; Component (thermodynamics); Software engineering; Process (computing); Systems engineering; Software; Engineering; Operating system","score_opus":0.05245264204775712,"score_gpt":0.26167256071511175,"score_spread":0.20921991866735462,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1517245390","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025477283,0.00037558642,0.94784313,0.00022676655,0.000044913373,0.00053841074,0.0012498068,0.02248027,0.0017638042],"genre_scores_gemma":[0.11225346,0.00022379083,0.87533194,0.00026239664,0.000032107942,0.0006428516,0.0057179295,0.0018998343,0.0036356046],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9972427,0.0004581256,0.00030988987,0.00044036837,0.0013773481,0.00017155707],"domain_scores_gemma":[0.9880648,0.006427041,0.00088817795,0.0017942775,0.0025556872,0.00026995872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001757489,0.0019977011,0.0013318666,0.005135375,0.0008189603,0.0028921075,0.004472512,0.0018563883,0.003917463],"category_scores_gemma":[0.011315418,0.0012065684,0.002518987,0.003445156,0.0011292783,0.0033199813,0.001989969,0.0015612325,0.0021628735],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005974738,0.0007046367,0.0063349144,0.0011518301,0.0004034662,0.0014051392,0.00057116547,0.031809546,0.08540893,0.008331282,0.011276537,0.852005],"study_design_scores_gemma":[0.00029454535,0.00067018776,0.0050956374,0.00019853505,0.0004984963,0.0018795584,0.00041472638,0.7694448,0.17113312,0.028477846,0.0216821,0.00021051092],"about_ca_topic_score_codex":0.006973547,"about_ca_topic_score_gemma":0.012841049,"teacher_disagreement_score":0.006973547,"about_ca_system_score_codex":0.0010419025,"about_ca_system_score_gemma":0.002728694,"threshold_uncertainty_score":0.013865888},"labels":[],"label_agreement":null},{"id":"W1518403088","doi":"10.1007/3-540-36209-6_29","title":"Product and Process Metrics: A Software Engineering Measurement Expert System","year":2002,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software Engineering Process Group; Social software engineering; Personal software process; Software construction; Software measurement; Software engineering; Software development; Software system; Benchmarking; Software sizing; Verification and validation; Software; Computer science; Engineering; Systems engineering; Operating system","score_opus":0.0304681407841012,"score_gpt":0.24273348211354376,"score_spread":0.21226534132944255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1518403088","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008532109,0.00046142508,0.89922035,0.00047016566,0.00007805188,0.0004205775,0.0016278387,0.07994546,0.009243992],"genre_scores_gemma":[0.041271675,0.00055115856,0.9302222,0.00022069598,0.00008025044,0.00052038743,0.0053154924,0.008038453,0.013779668],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965197,0.00061640923,0.00039287918,0.00041917645,0.0019689074,0.00008294246],"domain_scores_gemma":[0.9882129,0.0051210136,0.0006729684,0.0014208585,0.004144609,0.00042763626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063756667,0.0012985349,0.0014916572,0.004186794,0.0004726307,0.0021828162,0.0016697381,0.00095113495,0.008988231],"category_scores_gemma":[0.01885221,0.0014401322,0.0005010436,0.0028415152,0.00029552457,0.0040780585,0.0016821752,0.0014916295,0.008159204],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011263703,0.0001745908,0.0024976502,0.00020353493,0.00007253131,0.00009014391,0.0003778759,0.004997938,0.00803792,0.0036839724,0.058960333,0.9207909],"study_design_scores_gemma":[0.00045974375,0.0005156249,0.026046276,0.0006556067,0.00053945475,0.0014734834,0.000376119,0.5242749,0.052903835,0.04124486,0.3511589,0.0003511866],"about_ca_topic_score_codex":0.00192131,"about_ca_topic_score_gemma":0.0029608053,"teacher_disagreement_score":0.008988231,"about_ca_system_score_codex":0.00071269117,"about_ca_system_score_gemma":0.0017836833,"threshold_uncertainty_score":0.03371817},"labels":[],"label_agreement":null},{"id":"W1519321440","doi":"","title":"Software Maintenance Maturity Model (SM mm ): the software maintenance process model: Research Articles","year":2005,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Capability Maturity Model Integration; Capability Maturity Model; LeanCMMI; Software maintenance; Software engineering; Scope (computer science); Computer science; Verification and validation; Software development; Systems engineering; Software construction; Software development process; Software; Process management; Engineering; Operations management; Operating system","score_opus":0.06422356596106492,"score_gpt":0.3667582124005909,"score_spread":0.302534646439526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1519321440","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035895877,0.065913126,0.7326261,0.049225554,0.0006894869,0.0008017165,0.0007047567,0.0018965557,0.11224686],"genre_scores_gemma":[0.5253154,0.040711433,0.4168617,0.0023836321,0.00061351334,0.0014200642,0.0017075111,0.0002301239,0.010756636],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9957398,0.0013817076,0.00041256618,0.00037543144,0.0018686614,0.00022177574],"domain_scores_gemma":[0.9905505,0.00386926,0.0013526093,0.0005304745,0.003297132,0.00039999394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006376748,0.00079329166,0.00055518723,0.003531217,0.00089509774,0.004928234,0.0015483851,0.00338707,0.0017920296],"category_scores_gemma":[0.017862035,0.00035558725,0.0006998772,0.0055521782,0.0014438721,0.008234735,0.0017000134,0.0022747784,0.0010823465],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038124243,0.00015153014,0.005350417,0.0015659974,0.000051524425,0.00016187118,0.0020177064,0.009876622,0.0017499636,0.5042379,0.016514592,0.45828372],"study_design_scores_gemma":[0.00005711311,0.00030970867,0.008342654,0.0054317294,0.00017041476,0.0010972198,0.0018973657,0.06358513,0.00368324,0.5590556,0.35622817,0.00014172343],"about_ca_topic_score_codex":0.0039981813,"about_ca_topic_score_gemma":0.0026051633,"teacher_disagreement_score":0.006376748,"about_ca_system_score_codex":0.004415254,"about_ca_system_score_gemma":0.00816252,"threshold_uncertainty_score":0.03372383},"labels":[],"label_agreement":null},{"id":"W1520637757","doi":"10.1007/3-540-36103-0_49","title":"A Formal Definition of Function Points for Automated Measurement of B Specifications","year":2002,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Function point; Computer science; Function (biology); Programming language; Software engineering; Software; Software development; Biology","score_opus":0.08049403041518137,"score_gpt":0.2598502023984091,"score_spread":0.17935617198322773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1520637757","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00159088,0.00006756405,0.9959824,0.000112878835,0.000029046636,0.00006973755,0.00007824721,0.00075582543,0.0013134371],"genre_scores_gemma":[0.09570531,0.00021376007,0.90068877,0.0002312735,0.00010369126,0.00040644887,0.00025891626,0.00048176857,0.0019101248],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99342316,0.0018399801,0.001026348,0.0010462984,0.0020766214,0.0005876153],"domain_scores_gemma":[0.98817074,0.0055303276,0.001061715,0.003018473,0.0018768436,0.00034188599],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005067424,0.0017699585,0.0012805455,0.0037374224,0.0015552752,0.007909249,0.0032854618,0.0040772706,0.0048965337],"category_scores_gemma":[0.016441574,0.0016391355,0.0019290552,0.0024704963,0.005442947,0.0069161546,0.0035905426,0.004536997,0.0021730878],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000086646876,0.0000811584,0.0006042885,0.00014639976,0.000025728543,0.00022736378,0.00045365313,0.0071661496,0.011476565,0.9205316,0.002303934,0.056896444],"study_design_scores_gemma":[0.00008694937,0.00018769996,0.00051211094,0.00024547084,0.000060462233,0.0007257949,0.000196846,0.13583246,0.034891855,0.7968696,0.030263847,0.00012697943],"about_ca_topic_score_codex":0.0020339536,"about_ca_topic_score_gemma":0.0012305434,"teacher_disagreement_score":0.007909249,"about_ca_system_score_codex":0.001455914,"about_ca_system_score_gemma":0.0019359383,"threshold_uncertainty_score":0.02679944},"labels":[],"label_agreement":null},{"id":"W1520957802","doi":"","title":"Proceedings of the 16th ACM SIGSOFT International Symposium on Foundations of software engineering","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Presentation (obstetrics); Software engineering; Computer science; Event (particle physics); Engineering management; Software; Software development; Software peer review; Engineering; Software construction; Medicine","score_opus":0.020751618399376826,"score_gpt":0.24885801267680774,"score_spread":0.2281063942774309,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1520957802","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016361492,0.05804151,0.5082063,0.04377661,0.09299151,0.001276816,0.005120442,0.011162973,0.26306233],"genre_scores_gemma":[0.061576024,0.041438796,0.18938395,0.008510195,0.015198675,0.0010276092,0.014765488,0.0046349107,0.66346437],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9956986,0.0010455175,0.00033537394,0.00060439575,0.0019816628,0.00033461105],"domain_scores_gemma":[0.989674,0.0028684326,0.00038431134,0.002003648,0.0033320524,0.0017376738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005597843,0.0019025684,0.0017969649,0.0021598602,0.0012633897,0.0056995344,0.0018141147,0.0034936971,0.116025314],"category_scores_gemma":[0.013055361,0.0010465976,0.0016415289,0.0013875379,0.0018292503,0.0051267776,0.0032442948,0.0067590536,0.047712576],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036248929,0.0003498018,0.0018014532,0.00040013227,0.00010216868,0.0003787612,0.00036232686,0.0021834609,0.003655045,0.015509021,0.51562965,0.45926568],"study_design_scores_gemma":[0.00005579943,0.00017554329,0.0011877344,0.00037119852,0.0000386441,0.00039538735,0.00011969152,0.0035007857,0.00106406,0.016619233,0.9764282,0.000043676973],"about_ca_topic_score_codex":0.0041943574,"about_ca_topic_score_gemma":0.0048834886,"teacher_disagreement_score":0.116025314,"about_ca_system_score_codex":0.0016648655,"about_ca_system_score_gemma":0.004286376,"threshold_uncertainty_score":0.38814336},"labels":[],"label_agreement":null},{"id":"W1523215431","doi":"10.5772/38280","title":"A Simulation Approach to Validate Models Derived from Observational Studies","year":2012,"lang":"en","type":"book-chapter","venue":"InTech eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Observational study; Computer science; Mathematics; Statistics","score_opus":0.24507841858036658,"score_gpt":0.3353757487902126,"score_spread":0.09029733020984601,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1523215431","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031666566,0.00019994142,0.9347763,0.00054043584,0.00022906593,0.0021494983,0.0017581822,0.0009342748,0.027745616],"genre_scores_gemma":[0.33746436,0.00034542932,0.6498478,0.00022641454,0.000046907226,0.0069399336,0.0013402068,0.00023927122,0.003549647],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9901439,0.0077224555,0.000516991,0.00053031463,0.00086799444,0.0002183034],"domain_scores_gemma":[0.89706796,0.09055263,0.0017794344,0.005671414,0.0046068053,0.00032172017],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022208622,0.0009748492,0.0009874932,0.0022583555,0.001054137,0.0020349429,0.00308649,0.0017553638,0.016128195],"category_scores_gemma":[0.07949779,0.00071363145,0.0020239737,0.0018058814,0.0010804753,0.0023486295,0.00212637,0.0028907354,0.0012568556],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044362596,0.0004546145,0.0065728608,0.00074211013,0.00041629755,0.0002420288,0.0011619942,0.60593194,0.0009875358,0.33881584,0.005030804,0.039200407],"study_design_scores_gemma":[0.00026358853,0.00039471252,0.00078983355,0.000265685,0.00010729547,0.00005880411,0.0004295412,0.8978675,0.0011455055,0.08401112,0.014613673,0.00005277449],"about_ca_topic_score_codex":0.008161091,"about_ca_topic_score_gemma":0.005899212,"teacher_disagreement_score":0.022208622,"about_ca_system_score_codex":0.0020994714,"about_ca_system_score_gemma":0.004082025,"threshold_uncertainty_score":0.11745179},"labels":[],"label_agreement":null},{"id":"W1523392525","doi":"10.1007/978-3-540-78921-5_20","title":"Using FCA to Suggest Refactorings to Correct Design Defects","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Programming language; Artificial intelligence; Software engineering","score_opus":0.05870075889914268,"score_gpt":0.2911270673122262,"score_spread":0.23242630841308354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1523392525","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.105872296,0.0012272889,0.8447863,0.0021042353,0.00049305387,0.00052842806,0.0017013985,0.025259523,0.018027501],"genre_scores_gemma":[0.30144656,0.0002761322,0.6891403,0.0003141415,0.000084868254,0.00014572912,0.0018784296,0.0008173365,0.005896516],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985045,0.0002900728,0.00009758708,0.00032524497,0.00066542166,0.000117150965],"domain_scores_gemma":[0.98638606,0.007470533,0.0010149077,0.0017470056,0.0031436882,0.00023777691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019213348,0.0018114081,0.0007093303,0.0042111985,0.0010803095,0.0014367658,0.0020325417,0.002046391,0.009735971],"category_scores_gemma":[0.021020256,0.0005790027,0.0015221208,0.0012552381,0.00059715804,0.0019162968,0.00075804093,0.001579852,0.0015686181],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006025356,0.00038524784,0.011155248,0.0006943719,0.00015473318,0.00091805647,0.0007093158,0.060908448,0.021188613,0.00985066,0.023669994,0.8697628],"study_design_scores_gemma":[0.00012426921,0.0002571352,0.0027179115,0.00037186468,0.00030797604,0.0004909925,0.0003088864,0.9339781,0.025891382,0.017761523,0.017719131,0.000070940485],"about_ca_topic_score_codex":0.016157791,"about_ca_topic_score_gemma":0.035578698,"teacher_disagreement_score":0.016157791,"about_ca_system_score_codex":0.0012596155,"about_ca_system_score_gemma":0.0027431247,"threshold_uncertainty_score":0.032570064},"labels":[],"label_agreement":null},{"id":"W1524317667","doi":"10.1007/11767718_46","title":"Art and Science of System Release Planning","year":2006,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Intuition; Negotiation; Management science; Plan (archaeology); Decision support system; Operations research; Knowledge management; Artificial intelligence; Engineering","score_opus":0.014661575076145503,"score_gpt":0.2477978904997026,"score_spread":0.23313631542355712,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1524317667","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036557703,0.04810369,0.46835187,0.01858355,0.002343069,0.000100472294,0.00028665602,0.0005630712,0.45801196],"genre_scores_gemma":[0.3229269,0.09649629,0.35289037,0.0055404594,0.005871818,0.0005359384,0.00054195523,0.00068597327,0.21451029],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988689,0.0003291192,0.000062226274,0.00018471832,0.00047333227,0.00008165289],"domain_scores_gemma":[0.99779,0.0015506865,0.00008876286,0.00033153984,0.00017008898,0.00006897646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016762615,0.0008410448,0.0005946105,0.0013301743,0.0012876379,0.004681378,0.0015153,0.0015491423,0.012097467],"category_scores_gemma":[0.004315339,0.00081626925,0.00069107703,0.001995984,0.007050867,0.006092276,0.0011329639,0.0048527243,0.0023119266],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012617761,0.000014796534,0.00006021492,0.0001465954,0.0000061292067,0.00003499356,0.00021803356,0.0041772993,0.00020164308,0.93494874,0.014515333,0.045663647],"study_design_scores_gemma":[0.000007572061,0.000015293796,0.00010812525,0.0001323008,0.000007245985,0.00012631372,0.00006378057,0.005802654,0.00026133354,0.84074676,0.15271527,0.000013347859],"about_ca_topic_score_codex":0.0027152812,"about_ca_topic_score_gemma":0.0026065016,"teacher_disagreement_score":0.012097467,"about_ca_system_score_codex":0.0022478728,"about_ca_system_score_gemma":0.0025566819,"threshold_uncertainty_score":0.040470123},"labels":[],"label_agreement":null},{"id":"W1525022364","doi":"10.1109/iwpse.2003.1231214","title":"The chaos of software development","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software development; Suite; Software engineering; Source code; Package development process; Software; Software project management; Software evolution; Software development process; Software construction; Data science; Operating system","score_opus":0.014836853700048249,"score_gpt":0.24335710023699209,"score_spread":0.22852024653694383,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1525022364","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1650097,0.025151948,0.646592,0.053114995,0.0014848316,0.00018258087,0.0006053241,0.0008246869,0.10703396],"genre_scores_gemma":[0.94493157,0.0072253724,0.03754215,0.0015359494,0.001228445,0.00027060052,0.00014347192,0.00012974642,0.0069929133],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99753296,0.00091306266,0.00011647468,0.00044243218,0.00074063597,0.00025454798],"domain_scores_gemma":[0.98606384,0.008689409,0.0016027427,0.001925319,0.0011024884,0.0006161728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023871877,0.00041074303,0.00073448924,0.0025327227,0.002010565,0.0041226977,0.0007937961,0.0015607629,0.00223991],"category_scores_gemma":[0.014364927,0.00035362158,0.000857229,0.00149058,0.009270395,0.007790584,0.0037663886,0.0021478352,0.00044101343],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004262194,0.00001608714,0.0038291877,0.00018869658,0.00004681426,0.00020792245,0.0018123942,0.016762475,0.0010043519,0.94095224,0.0034341726,0.031703044],"study_design_scores_gemma":[0.000012500054,0.000039631323,0.0017123214,0.00005441403,0.000013188594,0.00012307357,0.00029016458,0.018251415,0.00027367842,0.959409,0.019791456,0.00002908715],"about_ca_topic_score_codex":0.0017843047,"about_ca_topic_score_gemma":0.0010487549,"teacher_disagreement_score":0.0041226977,"about_ca_system_score_codex":0.0020603647,"about_ca_system_score_gemma":0.0017238535,"threshold_uncertainty_score":0.014949083},"labels":[],"label_agreement":null},{"id":"W1526576188","doi":"10.1007/978-3-540-31797-5_22","title":"Managed Architecture of Existing Code as a Practical Transition Towards MDA","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Zymeworks (Canada)","funders":"","keywords":"Code refactoring; Architecture; Computer science; Container (type theory); Abstraction; Code (set theory); Software engineering; Programming language; Engineering","score_opus":0.0394487252914162,"score_gpt":0.314096621160483,"score_spread":0.2746478958690668,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1526576188","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08346361,0.00009819711,0.88224196,0.0010051357,0.00008412118,0.00018990511,0.00006952905,0.00852184,0.02432569],"genre_scores_gemma":[0.48942256,0.00012012664,0.49329796,0.00019338375,0.000024640382,0.00015666451,0.00023104755,0.0013975933,0.015156015],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99841857,0.0006256514,0.00009739934,0.00017958338,0.0005491443,0.0001296901],"domain_scores_gemma":[0.9958262,0.001004036,0.0001237653,0.0024511104,0.0004030961,0.0001916419],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002749869,0.0002450459,0.00023427336,0.00067848037,0.0006100516,0.0030859415,0.0015475316,0.0011005219,0.0049378285],"category_scores_gemma":[0.006277165,0.00052147196,0.0005265433,0.00043532817,0.0011492426,0.004146775,0.002942039,0.0020429087,0.0009918351],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026712264,0.00035865908,0.0027498584,0.0002449486,0.00006632761,0.0011846324,0.008591682,0.017529754,0.061636936,0.55029863,0.008013817,0.3490577],"study_design_scores_gemma":[0.00015562115,0.0006096212,0.0030733664,0.00043086187,0.00022377328,0.0017828235,0.0018301148,0.35313043,0.08407258,0.33179197,0.22277294,0.00012587126],"about_ca_topic_score_codex":0.0008543566,"about_ca_topic_score_gemma":0.001525888,"teacher_disagreement_score":0.0049378285,"about_ca_system_score_codex":0.00069081364,"about_ca_system_score_gemma":0.0013292498,"threshold_uncertainty_score":0.016518652},"labels":[],"label_agreement":null},{"id":"W1526698614","doi":"10.1002/9780470606834.app2","title":"Appendix B: Glossary of Terms in Software Measurement","year":2010,"lang":"en","type":"other","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; École de Technologie Supérieure","funders":"","keywords":"Glossary; Appendix; Computer science; Software; Programming language; Linguistics; Philosophy; Biology","score_opus":0.021379694645305584,"score_gpt":0.2540536733886974,"score_spread":0.2326739787433918,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1526698614","genre_codex":"other","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011551952,0.014847577,0.043305475,0.003695017,0.004226937,0.0013446071,0.43254113,0.0055364342,0.4933477],"genre_scores_gemma":[0.017972315,0.032104485,0.09348392,0.0067150914,0.0027890126,0.0042624017,0.48260605,0.014961804,0.345105],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99833244,0.00034597603,0.00039230846,0.00020547332,0.0006535798,0.00007020977],"domain_scores_gemma":[0.99214447,0.0039458717,0.00056309154,0.0006125816,0.0025653464,0.0001686338],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0014163373,0.001690511,0.0011681186,0.009244972,0.0012089463,0.0035247211,0.0013693332,0.0013404755,0.4181262],"category_scores_gemma":[0.013446689,0.00075646857,0.00063890335,0.01606121,0.0007076572,0.0054369643,0.0017559355,0.002417864,0.30260167],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027607335,0.000018997582,0.00013729124,0.0012868268,0.000005764306,0.00005032627,0.00016875117,0.00019733809,0.0002556403,0.015542557,0.9416051,0.04070383],"study_design_scores_gemma":[0.000006433349,0.0000044312683,0.0002450184,0.00048768526,0.0000035411656,0.00007309832,0.00005714062,0.00009571872,0.000089377296,0.003887962,0.9950375,0.00001192144],"about_ca_topic_score_codex":0.0072036097,"about_ca_topic_score_gemma":0.007934364,"teacher_disagreement_score":0.4181262,"about_ca_system_score_codex":0.0018883476,"about_ca_system_score_gemma":0.0018813565,"threshold_uncertainty_score":0.8299724},"labels":[],"label_agreement":null},{"id":"W1527416500","doi":"","title":"Relating Requirements to Implementation via Topic Analysis","year":2012,"lang":"en","type":"article","venue":"International Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Traceability; Documentation; Relevance (law); Topic model; Perception; Control (management); Requirements traceability; Requirements analysis; Requirements engineering; Software; Software engineering; Information retrieval; Artificial intelligence; Requirement","score_opus":0.060189474228836,"score_gpt":0.3648502363919053,"score_spread":0.30466076216306925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1527416500","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30306876,0.00050675974,0.68538356,0.00076214416,0.000051504867,0.0009684135,0.0013952085,0.00093660573,0.00692706],"genre_scores_gemma":[0.83216625,0.00025394387,0.16279346,0.000088581844,0.000054419383,0.0011241005,0.0022786635,0.000173436,0.0010669706],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98132,0.011820101,0.0012760037,0.0026028662,0.002434707,0.00054619426],"domain_scores_gemma":[0.8232012,0.14817609,0.009572992,0.007818277,0.010520365,0.0007110693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017883454,0.00079967204,0.0006370528,0.0074333106,0.0011835244,0.0040987944,0.0011649766,0.0012004,0.0022866575],"category_scores_gemma":[0.10238749,0.00062487746,0.0015775611,0.0056695854,0.0012700926,0.0045969216,0.0023712104,0.0020249556,0.0006877322],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014224765,0.0006642773,0.2671163,0.0020910667,0.0006832479,0.00050613034,0.06621384,0.0595178,0.016116587,0.048350677,0.006220521,0.5310971],"study_design_scores_gemma":[0.00013912279,0.0005877219,0.18447037,0.0005097653,0.00044154684,0.0006769878,0.02455922,0.6741496,0.012759029,0.076481394,0.024878368,0.00034690445],"about_ca_topic_score_codex":0.0075468663,"about_ca_topic_score_gemma":0.0050563463,"teacher_disagreement_score":0.017883454,"about_ca_system_score_codex":0.0023936909,"about_ca_system_score_gemma":0.0016191879,"threshold_uncertainty_score":0.09457785},"labels":[],"label_agreement":null},{"id":"W1527607785","doi":"10.1109/coginf.2003.1225966","title":"A cognitive complexity metric based on category learning","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Program comprehension; Comprehension; Computer science; Process (computing); Identifier; Set (abstract data type); Software; Software development; Metric (unit); Cognition; Software maintenance; Artificial intelligence; Software engineering; Software system; Human–computer interaction; Programming language; Engineering; Psychology","score_opus":0.05057610787086489,"score_gpt":0.29897688489369945,"score_spread":0.24840077702283456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1527607785","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24102083,0.002066438,0.6832986,0.0011460903,0.00031012285,0.0009804564,0.0031150726,0.0012260602,0.06683633],"genre_scores_gemma":[0.79387796,0.00055723666,0.1981106,0.00017072992,0.00015188698,0.0010518822,0.0020595149,0.00016362507,0.003856545],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99339175,0.0014345775,0.0005816306,0.00080562005,0.003433554,0.00035285],"domain_scores_gemma":[0.95750946,0.02869837,0.0038625605,0.0031356397,0.0052599995,0.0015340879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047605378,0.0008263785,0.00071573054,0.00898458,0.0009814093,0.0031921219,0.0013253391,0.0013470363,0.005594215],"category_scores_gemma":[0.052356873,0.00021227321,0.0010467322,0.0046686223,0.0026908554,0.006498048,0.0030073437,0.0015165443,0.00068870815],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075274636,0.0008277865,0.12294975,0.00086062547,0.000502648,0.0002620935,0.0020450724,0.041532136,0.0083407285,0.2603393,0.013002851,0.5485842],"study_design_scores_gemma":[0.000102606034,0.0019672958,0.13756782,0.00023555038,0.00018929514,0.0013381145,0.0014105594,0.20296822,0.008224703,0.61666673,0.028973088,0.00035604645],"about_ca_topic_score_codex":0.0025132585,"about_ca_topic_score_gemma":0.0026187755,"teacher_disagreement_score":0.00898458,"about_ca_system_score_codex":0.0026693041,"about_ca_system_score_gemma":0.0012431068,"threshold_uncertainty_score":0.025176406},"labels":[],"label_agreement":null},{"id":"W1530443667","doi":"10.1007/3-540-44839-x_76","title":"Relationships Between Selected Software Measures and Latent Bug-Density: Guidelines for Improving Quality","year":2003,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Carleton University","funders":"","keywords":"Computer science; Software quality; Quality (philosophy); Software development; Software engineering; Software bug; Software metric; Software; Work (physics); Software quality analyst; Verification and validation; Software peer review; Software construction; Data science; Engineering; Operations management; Programming language","score_opus":0.11930376333376365,"score_gpt":0.3259558737648782,"score_spread":0.2066521104311146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1530443667","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.122650445,0.010335526,0.8428535,0.005282808,0.0002923252,0.0011650809,0.0048756655,0.0053069806,0.0072376058],"genre_scores_gemma":[0.41327018,0.0017618524,0.57852507,0.00050965603,0.00017220694,0.0012156272,0.0024233295,0.0006146626,0.0015074624],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9799136,0.012130189,0.0024086542,0.0016076603,0.0036318677,0.00030808695],"domain_scores_gemma":[0.5746707,0.3417402,0.030618487,0.016359808,0.034474712,0.0021361469],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.039643068,0.002075643,0.0024895854,0.009286223,0.0011368684,0.0053461,0.0040366338,0.0023172572,0.0058586355],"category_scores_gemma":[0.29229856,0.0012666526,0.002059353,0.011257232,0.0017529356,0.0066854055,0.0017593874,0.0037114574,0.0017464277],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068966305,0.00093083066,0.43173954,0.002954799,0.002041575,0.00029839922,0.0021918428,0.015436837,0.0038911619,0.01660452,0.030585341,0.49263543],"study_design_scores_gemma":[0.0009019156,0.0029247827,0.58851904,0.0036219084,0.0070706075,0.0013781573,0.003266092,0.19884035,0.011307741,0.15993391,0.02161664,0.0006189819],"about_ca_topic_score_codex":0.0035101653,"about_ca_topic_score_gemma":0.007889025,"teacher_disagreement_score":0.039643068,"about_ca_system_score_codex":0.0013500134,"about_ca_system_score_gemma":0.002656394,"threshold_uncertainty_score":0.20965505},"labels":[],"label_agreement":null},{"id":"W1532113059","doi":"10.1007/978-3-540-69858-6_28","title":"Using Linguistic Knowledge to Classify Non-functional Requirements in SRS documents","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software requirements; Software engineering; Non-functional requirement; Software requirements specification; Software; Software development; Functional requirement; Classifier (UML); Natural language; Process (computing); Artificial intelligence; Natural language processing; Software construction; Programming language","score_opus":0.06269338556165842,"score_gpt":0.3261659636924095,"score_spread":0.2634725781307511,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1532113059","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6281047,0.001344096,0.32826746,0.0014815355,0.00012969367,0.0007960658,0.0046632322,0.0043287934,0.030884521],"genre_scores_gemma":[0.74655974,0.00043599147,0.24035028,0.00016676346,0.000063375766,0.0002620772,0.008001976,0.00026523173,0.0038945586],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963993,0.00082923856,0.000645981,0.00044227994,0.0015105244,0.00017271472],"domain_scores_gemma":[0.982418,0.0122636985,0.0017349473,0.0009259001,0.0024335152,0.00022392278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026264188,0.00044279467,0.00040520413,0.010163775,0.0008950061,0.00299729,0.00090507005,0.0010111261,0.0018470823],"category_scores_gemma":[0.014268353,0.00035156778,0.00080518273,0.003437243,0.00095268263,0.003223866,0.0010259559,0.0011230123,0.0011164816],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004660826,0.00057127693,0.045586053,0.0014958717,0.00017356593,0.0018028897,0.010515817,0.011164764,0.05794104,0.012840427,0.007407157,0.850035],"study_design_scores_gemma":[0.00029257403,0.0011179574,0.17751117,0.002616299,0.001460875,0.0056246794,0.023190688,0.50022364,0.10841374,0.071493804,0.10754863,0.0005058715],"about_ca_topic_score_codex":0.0067395293,"about_ca_topic_score_gemma":0.011393128,"teacher_disagreement_score":0.010163775,"about_ca_system_score_codex":0.0011947209,"about_ca_system_score_gemma":0.0020343764,"threshold_uncertainty_score":0.013889968},"labels":[],"label_agreement":null},{"id":"W1532805649","doi":"10.1002/spe.2134","title":"Validating pragmatic reuse tasks by leveraging existing test suites","year":2012,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Computer science; Correctness; Software engineering; Test suite; Test (biology); Code (set theory); Task (project management); Code reuse; Software; Test case; Programming language; Systems engineering; Machine learning; Engineering","score_opus":0.033974894645382156,"score_gpt":0.32588452441226484,"score_spread":0.2919096297668827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1532805649","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44090408,0.00030417222,0.54289013,0.00041206018,0.000089094385,0.0007962607,0.00028644426,0.010973012,0.0033447412],"genre_scores_gemma":[0.60440737,0.00016379243,0.39226037,0.00016924675,0.000028427534,0.00037157646,0.0010733121,0.0008518458,0.0006740442],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9540493,0.022030756,0.0047900053,0.0041883206,0.013460236,0.0014814304],"domain_scores_gemma":[0.7427518,0.16218735,0.017517183,0.05497807,0.020833945,0.0017316809],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023574343,0.0018297745,0.0011056109,0.003319969,0.00067248906,0.0025719407,0.0035378272,0.00188561,0.0016683524],"category_scores_gemma":[0.14559034,0.000914175,0.0014995108,0.0011567192,0.002087854,0.002769984,0.002876586,0.0019224361,0.0009079251],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009007201,0.0035278087,0.059126522,0.0015834367,0.0006798279,0.0022733225,0.003798576,0.13424921,0.17264232,0.0111418255,0.003777787,0.6062986],"study_design_scores_gemma":[0.00059201126,0.002815837,0.024216786,0.0006862694,0.00037274975,0.001900624,0.0009675661,0.7330743,0.20334984,0.017066248,0.014607396,0.0003503213],"about_ca_topic_score_codex":0.0021639832,"about_ca_topic_score_gemma":0.0027679156,"teacher_disagreement_score":0.023574343,"about_ca_system_score_codex":0.0012834668,"about_ca_system_score_gemma":0.0029245259,"threshold_uncertainty_score":0.1246745},"labels":[],"label_agreement":null},{"id":"W1534933129","doi":"10.1007/978-3-642-05415-0_12","title":"Functional Size Measurement Quality Challenges for Inexperienced Measurers","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Computer science; Software; Quality (philosophy); Engineering drawing; Algorithm; Data mining; Programming language; Engineering","score_opus":0.10274443870019853,"score_gpt":0.303481252921621,"score_spread":0.20073681422142245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1534933129","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3378954,0.009994751,0.5807449,0.025129082,0.00084834377,0.00022893216,0.0007334434,0.005899662,0.038525477],"genre_scores_gemma":[0.7655327,0.0020079087,0.21554963,0.0017585742,0.00061850593,0.00020467672,0.0005600507,0.001069236,0.012698757],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9677308,0.0064836848,0.0018625206,0.0027140821,0.020471552,0.0007373865],"domain_scores_gemma":[0.7872671,0.12481196,0.01045399,0.016899366,0.056095254,0.0044722524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029212952,0.00063643104,0.001037097,0.0024328262,0.0008956925,0.0042110546,0.0038052588,0.0011754704,0.0038499539],"category_scores_gemma":[0.14993554,0.0006253056,0.00040438262,0.001997569,0.0014521056,0.0048951562,0.0030839383,0.001470084,0.0015906753],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024554602,0.00014920185,0.03140148,0.0004475104,0.00007225611,0.00018968048,0.004067712,0.0038248983,0.010068361,0.013693341,0.014838494,0.9210015],"study_design_scores_gemma":[0.0002903616,0.0029668165,0.26333618,0.0018316652,0.00038226583,0.0065879757,0.01968046,0.2131024,0.09954877,0.2085824,0.18303363,0.00065698876],"about_ca_topic_score_codex":0.004079705,"about_ca_topic_score_gemma":0.0049292212,"teacher_disagreement_score":0.029212952,"about_ca_system_score_codex":0.0015942621,"about_ca_system_score_gemma":0.0018069458,"threshold_uncertainty_score":0.1544947},"labels":[],"label_agreement":null},{"id":"W1536065123","doi":"10.1109/icsm.2004.1357868","title":"Tools for extracting software structure from compiled programs","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Executable; Computer science; Suite; Software; Object (grammar); Software system; Software engineering; Object-oriented programming; Programming language; Artificial intelligence","score_opus":0.039137212980672326,"score_gpt":0.2851299881034815,"score_spread":0.24599277512280918,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1536065123","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070380014,0.0003376757,0.9672396,0.00006143928,0.000028577202,0.00009291282,0.00033024143,0.023995463,0.00087609765],"genre_scores_gemma":[0.043679237,0.00085251534,0.9479959,0.000056511504,0.000030435067,0.00023308586,0.0027385815,0.0031433285,0.0012704689],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977787,0.00034261678,0.00031615546,0.00041810804,0.0010028733,0.0001416002],"domain_scores_gemma":[0.9888728,0.0063302577,0.0011608384,0.0020960714,0.0013859194,0.00015414826],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017833199,0.0020697499,0.0015788076,0.0063073854,0.0011213965,0.002501612,0.0020558175,0.0010646244,0.0034023297],"category_scores_gemma":[0.021304347,0.0016884201,0.0020669627,0.0034509727,0.00092387974,0.0045968387,0.0019593225,0.0023451129,0.0029231748],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002294276,0.00030973423,0.0056160907,0.0016996778,0.00032802648,0.0008571576,0.0014467686,0.022024905,0.050898485,0.026288314,0.007851701,0.88244975],"study_design_scores_gemma":[0.00033732023,0.00047295922,0.008253502,0.0012476661,0.0006414425,0.0036291115,0.0007876616,0.44805008,0.33079055,0.08822298,0.11710251,0.00046415644],"about_ca_topic_score_codex":0.001607139,"about_ca_topic_score_gemma":0.0024263763,"teacher_disagreement_score":0.0063073854,"about_ca_system_score_codex":0.0005935288,"about_ca_system_score_gemma":0.0021688535,"threshold_uncertainty_score":0.011381924},"labels":[],"label_agreement":null},{"id":"W1536501326","doi":"10.1109/apsec.2004.35","title":"Automatic Detecting Code Cooperation","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Program comprehension; Software engineering; KPI-driven code analysis; Source code; Legacy system; Software maintenance; Static program analysis; Code (set theory); Software evolution; Legacy code; Software; Code review; Software development; Programming language; Software system; Software construction","score_opus":0.01855858642270239,"score_gpt":0.2747410288626309,"score_spread":0.2561824424399285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1536501326","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18619706,0.0005051661,0.7924233,0.00026335727,0.000059164733,0.00038561804,0.0006644008,0.014987376,0.0045146067],"genre_scores_gemma":[0.41959342,0.0002776425,0.57471544,0.00008387872,0.000029311375,0.00023028076,0.001808278,0.00079915905,0.0024625272],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949523,0.0010416016,0.00033106227,0.0012435793,0.0020853262,0.00034612874],"domain_scores_gemma":[0.98333514,0.0066489787,0.0028862352,0.0024433811,0.0042352704,0.0004510121],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002160875,0.00080611353,0.0008074663,0.007307211,0.0010464535,0.0015651255,0.0013539045,0.0012603614,0.0012835066],"category_scores_gemma":[0.015075827,0.00049489253,0.00066694,0.0024510801,0.00070182653,0.0022554048,0.0023069123,0.0009453398,0.0009541119],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041220788,0.00024913726,0.05371938,0.0006435812,0.0001083673,0.0010148657,0.0020195302,0.008234386,0.16217977,0.0121256,0.008515754,0.7507774],"study_design_scores_gemma":[0.000121755285,0.0003046503,0.041931503,0.00019361686,0.00020156406,0.0033068156,0.0014475237,0.6096123,0.27467585,0.023954755,0.044082128,0.0001675191],"about_ca_topic_score_codex":0.0023166232,"about_ca_topic_score_gemma":0.0029381688,"teacher_disagreement_score":0.007307211,"about_ca_system_score_codex":0.00070893526,"about_ca_system_score_gemma":0.0021910674,"threshold_uncertainty_score":0.011427939},"labels":[],"label_agreement":null},{"id":"W1536910550","doi":"10.1007/978-3-642-02047-6_10","title":"Academic Software Development Tools and Techniques","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Code refactoring; Software engineering; Context (archaeology); Object-oriented programming; Programming language; Software development; Focus (optics); Theme (computing); Software; World Wide Web","score_opus":0.025594859704889478,"score_gpt":0.274416134015549,"score_spread":0.2488212743106595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1536910550","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002536395,0.020269413,0.65956324,0.0019327158,0.0012063954,0.00018400956,0.0005918889,0.009767523,0.30394843],"genre_scores_gemma":[0.018425541,0.035152093,0.37248814,0.0005489646,0.0007234018,0.00032135803,0.00256616,0.0041231383,0.5656512],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988393,0.0001107746,0.00008924117,0.00011829515,0.00079631625,0.00004611787],"domain_scores_gemma":[0.99841475,0.0006732226,0.00007852341,0.0003238966,0.0004206241,0.00008892269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00087122124,0.0012941097,0.0007181862,0.0038581372,0.00070890324,0.0025592234,0.0017118239,0.00078056304,0.036356844],"category_scores_gemma":[0.0030981915,0.0011155407,0.0008460714,0.0059475186,0.00074602297,0.0036441947,0.001464336,0.0024621496,0.033950984],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011802903,0.000049982948,0.0001106136,0.00048921694,0.000012198907,0.000049970844,0.00025209595,0.0011858424,0.002389,0.085332714,0.088406876,0.8217097],"study_design_scores_gemma":[0.000011578594,0.000024311901,0.0003338616,0.00039247095,0.000021681435,0.00043740877,0.0000493691,0.0034194272,0.0034524503,0.06708447,0.92475325,0.000019545483],"about_ca_topic_score_codex":0.00045828603,"about_ca_topic_score_gemma":0.00088972936,"teacher_disagreement_score":0.036356844,"about_ca_system_score_codex":0.00087678793,"about_ca_system_score_gemma":0.0015978653,"threshold_uncertainty_score":0.12162572},"labels":[],"label_agreement":null},{"id":"W1537659679","doi":"10.24297/ijct.v13i5.2535","title":"An Analysis of the PROMISE and ISBSG Software Engineering Data Repositories","year":2014,"lang":"en","type":"article","venue":"INTERNATIONAL JOURNAL OF COMPUTERS & TECHNOLOGY","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Benchmarking; Computer science; Reusability; Software engineering; Data science; Software; Social software engineering; Software development; Software construction","score_opus":0.009101808855303768,"score_gpt":0.26729860595681165,"score_spread":0.2581967971015079,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1537659679","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75805163,0.002941792,0.053319454,0.01739237,0.0003171943,0.0013313281,0.14130387,0.008187835,0.017154474],"genre_scores_gemma":[0.6086745,0.0014226496,0.110872544,0.0011427518,0.0002940393,0.0019736018,0.26679462,0.001949226,0.006876034],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.94988817,0.017201118,0.0051089735,0.0031775632,0.023038108,0.0015860073],"domain_scores_gemma":[0.7589238,0.12399944,0.017231192,0.040319312,0.055964697,0.0035615996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04270231,0.0005799756,0.00089813786,0.019531759,0.0018632713,0.004854998,0.0028323182,0.002161346,0.0024033596],"category_scores_gemma":[0.15817407,0.00090604805,0.0015058231,0.029594015,0.0016163222,0.011055465,0.0051958933,0.0027520182,0.0016483944],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002176094,0.0014141282,0.42136523,0.005815804,0.0005054634,0.0018970559,0.021945093,0.009448543,0.01118047,0.050335515,0.14236835,0.33154824],"study_design_scores_gemma":[0.00020108087,0.0011539654,0.5424777,0.0016788099,0.00034307785,0.0029870763,0.020656738,0.049328756,0.018640323,0.011135245,0.35085818,0.00053893373],"about_ca_topic_score_codex":0.0050868713,"about_ca_topic_score_gemma":0.0075422493,"teacher_disagreement_score":0.04270231,"about_ca_system_score_codex":0.0024989638,"about_ca_system_score_gemma":0.0045506647,"threshold_uncertainty_score":0.22583407},"labels":[],"label_agreement":null},{"id":"W154166809","doi":"10.1007/978-1-84800-044-5_1","title":"Software Engineering Data Collection for Field Studies","year":2007,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":113,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; National Research Council Canada","funders":"","keywords":"Computer science; Software engineering; Software; Field (mathematics); Data collection; Data science; Software analytics; Software development; Software construction; Programming language","score_opus":0.14418180368235065,"score_gpt":0.34834739552624383,"score_spread":0.20416559184389318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W154166809","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005450436,0.0097526265,0.77577263,0.004368217,0.0012643231,0.003202637,0.029764175,0.010569277,0.15985562],"genre_scores_gemma":[0.032253824,0.011913913,0.80906403,0.0012490669,0.0006420044,0.0040445114,0.04253349,0.003065374,0.09523385],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9930494,0.0023310946,0.0008031668,0.0008694633,0.0027369175,0.00020988856],"domain_scores_gemma":[0.969467,0.012043355,0.00106768,0.009109127,0.0077347425,0.00057818356],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009651607,0.0013697937,0.0015437426,0.009129683,0.0017189329,0.004361678,0.0027721135,0.0009376022,0.032581877],"category_scores_gemma":[0.025890812,0.00094376487,0.00083034096,0.011287979,0.0012558001,0.005273249,0.0037388518,0.0018567921,0.028044475],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046112535,0.000078429075,0.0016910358,0.0008324733,0.00001990377,0.000043826924,0.0004983565,0.00036543497,0.0024058085,0.018899761,0.17783497,0.7972839],"study_design_scores_gemma":[0.00002804852,0.000060789796,0.007677427,0.0016344371,0.00005004383,0.00033459766,0.0014595259,0.0032758038,0.009329375,0.051795155,0.92428225,0.00007260802],"about_ca_topic_score_codex":0.0056118616,"about_ca_topic_score_gemma":0.008710519,"teacher_disagreement_score":0.032581877,"about_ca_system_score_codex":0.0017924878,"about_ca_system_score_gemma":0.0052099414,"threshold_uncertainty_score":0.108997226},"labels":[],"label_agreement":null},{"id":"W1542960140","doi":"10.1007/978-3-642-01853-4_5","title":"Communicating Domain Knowledge in Executable Acceptance Test Driven Development","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Executable; Computer science; Software engineering; Domain (mathematical analysis); Software development; Software; Test (biology); Knowledge management; Programming language","score_opus":0.016997622284350113,"score_gpt":0.260130751329196,"score_spread":0.24313312904484588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1542960140","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030025745,0.00017689026,0.9598914,0.000342399,0.000028909608,0.00012352005,0.000057753386,0.0020765213,0.0072767804],"genre_scores_gemma":[0.66364574,0.00028444335,0.33065957,0.00022627508,0.000046977548,0.00029144876,0.00037687455,0.000543384,0.003925365],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99376243,0.0030369058,0.0003654918,0.0004348646,0.0020399196,0.00036050408],"domain_scores_gemma":[0.9575475,0.035912212,0.0007947804,0.0038164468,0.0016624657,0.00026656283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044403877,0.00065263454,0.00060388516,0.0010695257,0.00069542264,0.002711841,0.0014513528,0.0021879203,0.003482213],"category_scores_gemma":[0.03543103,0.00086985005,0.00067088765,0.00079232955,0.0019293692,0.0047021485,0.0032152827,0.0022311881,0.0009401803],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063875655,0.0004642747,0.0034125233,0.0003869271,0.000085422784,0.0014570454,0.0027472745,0.15764466,0.015586682,0.17579989,0.0044015897,0.637375],"study_design_scores_gemma":[0.000048623857,0.000096541975,0.0005715553,0.00012922268,0.000049462542,0.00034736306,0.00018708351,0.79731697,0.021147905,0.17428675,0.0057746097,0.000043856253],"about_ca_topic_score_codex":0.0016889258,"about_ca_topic_score_gemma":0.0020221167,"teacher_disagreement_score":0.0044403877,"about_ca_system_score_codex":0.00062392716,"about_ca_system_score_gemma":0.0010962455,"threshold_uncertainty_score":0.023483276},"labels":[],"label_agreement":null},{"id":"W1543881559","doi":"10.1002/spe.2228","title":"A recommendation system for repairing violations detected by static architecture conformance checking","year":2013,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code refactoring; Conformance checking; Computer science; Software engineering; Architecture; Process (computing); Systems architecture; Process management; Engineering; Operations management; Programming language; Work in process; Business process; Software","score_opus":0.013843532244916478,"score_gpt":0.28333398793798403,"score_spread":0.26949045569306757,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1543881559","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022307463,0.00062885694,0.44478548,0.0011756226,0.00037836324,0.0019979542,0.0061059957,0.51508987,0.007530406],"genre_scores_gemma":[0.12028915,0.0009211389,0.8238632,0.0008810328,0.000184388,0.0012953925,0.020385513,0.012377455,0.019802732],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98914206,0.0025578076,0.0024705282,0.0016707911,0.0036813677,0.00047735168],"domain_scores_gemma":[0.9645466,0.011675688,0.003615156,0.008244595,0.010994749,0.0009230818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013959465,0.002410466,0.0014826045,0.008945057,0.0013677191,0.003084198,0.004155543,0.0032065515,0.018450111],"category_scores_gemma":[0.0459481,0.0014858149,0.0014184944,0.0027681605,0.0005636522,0.0033619583,0.0019519951,0.0016723895,0.0127261635],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015616842,0.0007119435,0.01564437,0.0012293305,0.00028524682,0.0015455057,0.0012654854,0.012591923,0.01965352,0.0051085516,0.14270467,0.79769766],"study_design_scores_gemma":[0.0012560497,0.0014126496,0.014852307,0.0017647516,0.0006617364,0.0022089628,0.0012360188,0.4213023,0.07795895,0.010594265,0.46564376,0.0011082196],"about_ca_topic_score_codex":0.014240922,"about_ca_topic_score_gemma":0.017754452,"teacher_disagreement_score":0.018450111,"about_ca_system_score_codex":0.001498065,"about_ca_system_score_gemma":0.0028527323,"threshold_uncertainty_score":0.0738256},"labels":[],"label_agreement":null},{"id":"W1544109115","doi":"10.1108/02656710510577224","title":"Object‐oriented software development antecedents that influence product bug density","year":2005,"lang":"en","type":"article","venue":"International Journal of Quality & Reliability Management","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Carleton University","funders":"","keywords":"Quality (philosophy); Computer science; Originality; Software quality; Product (mathematics); Software engineering; Software; Software development; Suite; Antecedent (behavioral psychology); New product development; Process management; Engineering; Psychology; Marketing; Business; Mathematics; Programming language","score_opus":0.02277272355202363,"score_gpt":0.3220412297818265,"score_spread":0.29926850622980283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1544109115","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9967097,0.00013936979,0.0008278251,0.0001617256,0.000005275676,0.000029227207,0.00003453301,0.000016772508,0.0020756742],"genre_scores_gemma":[0.99899536,0.00006916023,0.00060514535,0.000011967992,0.0000031827028,0.000012368603,0.000031861062,0.00000542078,0.00026545327],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9971602,0.0012297798,0.0002039097,0.0003167735,0.0008569435,0.00023247632],"domain_scores_gemma":[0.9058697,0.063152045,0.018737664,0.0027118665,0.0048413496,0.0046874015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030741699,0.0002964513,0.00027595786,0.0017650435,0.0006138614,0.0017223408,0.00044556696,0.00052814174,0.0059365286],"category_scores_gemma":[0.048858043,0.00042058763,0.0003301134,0.0010001004,0.00094236,0.0008086363,0.0012535679,0.001139845,0.00030850712],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010223617,0.0004394662,0.9855428,0.00005383379,0.000048126283,0.00015856228,0.00089092203,0.0005247354,0.00062636327,0.0011405841,0.0001717529,0.010300639],"study_design_scores_gemma":[0.000013596375,0.00016687988,0.9949433,0.00005438617,0.0000682227,0.00011353227,0.0010358206,0.0019482567,0.0004215153,0.00069919054,0.00052670954,0.0000085697375],"about_ca_topic_score_codex":0.0043019974,"about_ca_topic_score_gemma":0.005051054,"teacher_disagreement_score":0.0059365286,"about_ca_system_score_codex":0.000934094,"about_ca_system_score_gemma":0.001754389,"threshold_uncertainty_score":0.019859672},"labels":[],"label_agreement":null},{"id":"W1545579570","doi":"10.1002/smr.1636","title":"On the evolution of Lehman's Laws","year":2013,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Waterloo","funders":"","keywords":"Software evolution; Honor; Software; Software development; Computer science; Software engineering; Software construction; Programming language","score_opus":0.010730859224768127,"score_gpt":0.24545486193871513,"score_spread":0.234724002713947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1545579570","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08301984,0.015875561,0.27919826,0.024412416,0.0011457793,0.00008205866,0.0002014603,0.00032236488,0.5957423],"genre_scores_gemma":[0.903808,0.004859653,0.041420873,0.0025325457,0.00092199125,0.00012781825,0.00015763665,0.00016160328,0.046009995],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9986817,0.0005616638,0.000055556804,0.00017708342,0.00039623425,0.0001276908],"domain_scores_gemma":[0.9952434,0.0031382984,0.0002566214,0.00043789405,0.00075063953,0.0001731993],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002839469,0.0002474798,0.00042925795,0.0012968039,0.00151802,0.0021002076,0.00076864846,0.0011393353,0.0061237244],"category_scores_gemma":[0.0131402295,0.00023008842,0.00051690405,0.00078509934,0.005388737,0.00405321,0.0014640257,0.0026164935,0.00078269554],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00000312556,0.0000035253595,0.000089919085,0.000008458043,0.000001679554,0.000023941622,0.00017405357,0.0006684543,0.000034888377,0.9943954,0.0012600598,0.003336569],"study_design_scores_gemma":[0.000003384622,0.0000038677863,0.000117373704,0.000018161027,0.0000014133176,0.000022781833,0.000038225735,0.0031741897,0.000052629362,0.988138,0.008425642,0.000004421099],"about_ca_topic_score_codex":0.0031411496,"about_ca_topic_score_gemma":0.0016921874,"teacher_disagreement_score":0.0061237244,"about_ca_system_score_codex":0.0023214763,"about_ca_system_score_gemma":0.000992528,"threshold_uncertainty_score":0.020485878},"labels":[],"label_agreement":null},{"id":"W1546188910","doi":"","title":"Investigation of the metrology concepts in ISO 9126 on software product quality evaluation","year":2006,"lang":"en","type":"article","venue":"Annual Conference on Computers","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Metrology; Software; Terminology; Software measurement; Measurement uncertainty; Systems engineering; Quality assurance; Documentation; Computer science; Software quality; Quality (philosophy); Reliability engineering; Software engineering; Engineering; Software development; Mathematics; Programming language; Statistics; Operations management","score_opus":0.06600007355852455,"score_gpt":0.3392908951000278,"score_spread":0.27329082154150325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1546188910","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059064504,0.030136889,0.7506917,0.013668281,0.0020106176,0.0022126604,0.0003467381,0.00047008292,0.14139852],"genre_scores_gemma":[0.42041078,0.008469252,0.5555572,0.0031892557,0.0007474106,0.0021671525,0.00055152504,0.00037639035,0.008531041],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8716817,0.045152087,0.014008718,0.0053665806,0.062116828,0.0016740467],"domain_scores_gemma":[0.8731761,0.06895317,0.0125578735,0.006683664,0.037944797,0.0006843431],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06959539,0.0015785242,0.0013907262,0.012712099,0.0033289476,0.013146234,0.0031036984,0.004304455,0.0020961342],"category_scores_gemma":[0.1437004,0.001277579,0.0013432833,0.011842625,0.012521392,0.013860201,0.0043016155,0.0069520404,0.0007198852],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007393423,0.00010801828,0.0027123054,0.0009377925,0.000025580463,0.00023421884,0.007212997,0.0029950205,0.0033297227,0.8646658,0.0034353605,0.114269264],"study_design_scores_gemma":[0.00009100156,0.001071058,0.013869558,0.008685144,0.00012209507,0.002900773,0.011800989,0.026859235,0.013447642,0.41498274,0.50574225,0.00042742968],"about_ca_topic_score_codex":0.0042542475,"about_ca_topic_score_gemma":0.004598565,"teacher_disagreement_score":0.06959539,"about_ca_system_score_codex":0.008173935,"about_ca_system_score_gemma":0.016454017,"threshold_uncertainty_score":0.36806},"labels":[],"label_agreement":null},{"id":"W1548497223","doi":"10.5772/5936","title":"How Do Programmers Think?","year":2008,"lang":"en","type":"book-chapter","venue":"InTech eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Saint Mary's University","funders":"","keywords":"Computer science; Regular expression; Programming language; Alternation (linguistics); Concatenation (mathematics); Theoretical computer science; Expression (computer science); Compiler; Context (archaeology); Closure (psychology); Perl; Operator (biology); Control flow; Mathematics; Arithmetic; Linguistics","score_opus":0.03054038191266909,"score_gpt":0.2501165253759059,"score_spread":0.2195761434632368,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1548497223","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012399372,0.0131557165,0.032912485,0.44904682,0.008040031,0.00011531145,0.00070249196,0.00172679,0.48190102],"genre_scores_gemma":[0.39211488,0.024669966,0.046281643,0.22733192,0.0052865124,0.00063525204,0.0015901768,0.0027500906,0.29933953],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99202055,0.004332322,0.00026263032,0.0010330462,0.001538532,0.00081292423],"domain_scores_gemma":[0.9866508,0.0051364247,0.0009189165,0.001299306,0.0035870112,0.0024074612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006737863,0.0007590435,0.00043010255,0.0012254133,0.0032633164,0.011481725,0.0015070258,0.003025153,0.02607176],"category_scores_gemma":[0.026083147,0.00046724675,0.00040146877,0.0015786453,0.008395239,0.016714513,0.0033607122,0.005985298,0.016694188],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046885805,0.000059540067,0.004529764,0.00031059794,0.000028272105,0.00023064141,0.03965375,0.00013533286,0.00030241182,0.2672934,0.55660075,0.1308087],"study_design_scores_gemma":[0.0000097172415,0.000014489891,0.00074217224,0.00046958134,0.000009282945,0.00025564813,0.019743782,0.00014649171,0.00013987151,0.09204684,0.88639724,0.000024918923],"about_ca_topic_score_codex":0.003243086,"about_ca_topic_score_gemma":0.0025326733,"teacher_disagreement_score":0.02607176,"about_ca_system_score_codex":0.002327066,"about_ca_system_score_gemma":0.0040502716,"threshold_uncertainty_score":0.0872187},"labels":[],"label_agreement":null},{"id":"W1549553848","doi":"10.1023/a:1009815306478","title":"Replicated Case Studies for Investigating Quality Factors in Object-Oriented Designs","year":2001,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":172,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Carleton University","funders":"","keywords":"Computer science; Cohesion (chemistry); Software engineering; Quality (philosophy); Software quality; Set (abstract data type); Software; Data science; Data mining; Programming language; Software development","score_opus":0.1714365202462394,"score_gpt":0.4017966966694224,"score_spread":0.23036017642318302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1549553848","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85220027,0.0010675165,0.13643506,0.0002688266,0.00007932439,0.0026695651,0.00020846067,0.0001720441,0.0068987603],"genre_scores_gemma":[0.8847176,0.00038568937,0.11228758,0.00007154132,0.000029008836,0.0015907215,0.00015007198,0.000030516461,0.0007373183],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9327002,0.05449029,0.002794588,0.0019036072,0.0075100306,0.00060136773],"domain_scores_gemma":[0.5909102,0.31444126,0.018235423,0.058225106,0.016257698,0.0019302604],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03770876,0.000669368,0.0007544509,0.0028995243,0.0016648599,0.0020272657,0.0031659368,0.0025314076,0.0033028196],"category_scores_gemma":[0.23253487,0.00077361707,0.0010804264,0.0024994335,0.0020305736,0.0036264067,0.002375607,0.0013830292,0.00032908586],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0069904695,0.024842946,0.34752685,0.0042994823,0.002162767,0.0066738348,0.049685966,0.03797113,0.043212358,0.09964473,0.0028753725,0.37411416],"study_design_scores_gemma":[0.008252302,0.06396508,0.2910159,0.002443034,0.0044408953,0.011118627,0.054018922,0.27219704,0.076450616,0.18501565,0.030254671,0.000827237],"about_ca_topic_score_codex":0.0024738442,"about_ca_topic_score_gemma":0.005591773,"teacher_disagreement_score":0.96229124,"about_ca_system_score_codex":0.0021167777,"about_ca_system_score_gemma":0.0021471914,"threshold_uncertainty_score":0.19942534},"labels":[],"label_agreement":null},{"id":"W154983637","doi":"","title":"Perspectives on Formal Methods in the Last 25 years","year":2006,"lang":"en","type":"article","venue":"School of Computing Science Technical Report Series","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Formal methods; Context (archaeology); Computer science; Field (mathematics); Subject (documents); Quarter (Canadian coin); Software engineering; Work (physics); Data science; Engineering; History; Library science; Archaeology; Mathematics","score_opus":0.017221555868217848,"score_gpt":0.3558945100244626,"score_spread":0.33867295415624477,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W154983637","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00272047,0.70917857,0.017536348,0.1625702,0.009716829,0.000013775086,0.00009930138,0.000096153264,0.098068364],"genre_scores_gemma":[0.15468673,0.6935698,0.02139594,0.055546314,0.028604787,0.00015883478,0.0002269084,0.00020852072,0.04560221],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9941135,0.0029856383,0.0003154941,0.0006640622,0.0015012095,0.0004200849],"domain_scores_gemma":[0.9897466,0.0076316427,0.00046161754,0.00046201746,0.0011544547,0.00054364145],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009865372,0.0010009339,0.00077083067,0.0040214695,0.0028507647,0.0068131406,0.0014213928,0.0046648327,0.008748523],"category_scores_gemma":[0.008812729,0.00033500887,0.0009757181,0.0031668728,0.018246476,0.010277149,0.0030835236,0.006180328,0.0016453367],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000021681708,0.00001878268,0.00010146057,0.0002589022,0.000008202009,0.00006148973,0.0011250043,0.0003607578,0.00010072652,0.96081305,0.014733004,0.022396894],"study_design_scores_gemma":[0.000009970878,0.000025335841,0.00013350701,0.0006067841,0.000003858808,0.000121471094,0.0007052889,0.00028002693,0.00006216728,0.37165514,0.6263803,0.000016146147],"about_ca_topic_score_codex":0.0035832177,"about_ca_topic_score_gemma":0.002909918,"teacher_disagreement_score":0.99013466,"about_ca_system_score_codex":0.008899465,"about_ca_system_score_gemma":0.0028504503,"threshold_uncertainty_score":0.06457043},"labels":[],"label_agreement":null},{"id":"W1550175302","doi":"","title":"Supporting maintenance of legacy software with data mining techniques","year":2000,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Software maintenance; Computer science; Backporting; Software construction; Software analytics; Software engineering; Software development; Legacy system; Software system; Package development process; Software; Operating system","score_opus":0.10423546137047178,"score_gpt":0.41444658521642247,"score_spread":0.3102111238459507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1550175302","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04799052,0.0010166988,0.9434146,0.0014878507,0.000051092975,0.00049041933,0.0010432089,0.0031057177,0.0013998647],"genre_scores_gemma":[0.17543082,0.00091446954,0.8199603,0.00017749387,0.00008406145,0.0004903283,0.0023968942,0.00008985769,0.00045576418],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943932,0.002387001,0.0007248946,0.00083849963,0.0015156054,0.00014064803],"domain_scores_gemma":[0.95175517,0.037228387,0.00244058,0.0047107963,0.0035002592,0.0003648072],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010988762,0.001141714,0.0014586083,0.006894204,0.0011508749,0.0032009522,0.002778316,0.0012463559,0.0007969915],"category_scores_gemma":[0.03958536,0.0008021765,0.0018491432,0.0061490946,0.00081010594,0.0039373836,0.0016063428,0.0019000951,0.00075407757],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002768625,0.0011396105,0.037481613,0.0013840129,0.0005184451,0.0003444693,0.0018570324,0.030649101,0.012084142,0.0044786246,0.0034087535,0.9063774],"study_design_scores_gemma":[0.00025491533,0.00096361793,0.02596034,0.001182043,0.00081880385,0.0014575126,0.0029587976,0.7921893,0.06200342,0.077342145,0.03458898,0.00028008295],"about_ca_topic_score_codex":0.002232692,"about_ca_topic_score_gemma":0.0032464867,"teacher_disagreement_score":0.010988762,"about_ca_system_score_codex":0.00067630614,"about_ca_system_score_gemma":0.0016546438,"threshold_uncertainty_score":0.058114767},"labels":[],"label_agreement":null},{"id":"W1550431237","doi":"10.1023/a:1018924724621","title":"Software assessment using metrics: A comparison across large C++ and Java systems","year":2000,"lang":"en","type":"article","venue":"Annals of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bell (Canada); Polytechnique Montréal","funders":"","keywords":"Computer science; Software engineering; Software quality; Software system; Software; Java; Software sizing; Programming language; Verification and validation; Software measurement; Software metric; Software construction; Software development; Engineering","score_opus":0.05841628753574674,"score_gpt":0.3613897420495466,"score_spread":0.30297345451379987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1550431237","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98571813,0.0013718687,0.007940943,0.00035442537,0.000062614345,0.00011387212,0.0002799052,0.00037300217,0.003785211],"genre_scores_gemma":[0.991779,0.0003154919,0.006806211,0.000056890527,0.00002309821,0.00005088684,0.00042341367,0.00010837042,0.00043663333],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9908845,0.0029104995,0.00086267333,0.0007009379,0.0043922663,0.0002491627],"domain_scores_gemma":[0.85763276,0.09523394,0.009502793,0.005101653,0.027980104,0.0045486926],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0111226,0.0005072835,0.0006154792,0.0061903666,0.0008298347,0.0021026896,0.0010888501,0.0007791273,0.0011370893],"category_scores_gemma":[0.10583729,0.00024786283,0.0005376424,0.0050191134,0.00070285297,0.0033549839,0.0015286031,0.0005799792,0.00035467028],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004591861,0.0014349518,0.41887215,0.0011414974,0.00064942386,0.00028300306,0.005233503,0.0060875504,0.01609281,0.0040990324,0.004644996,0.5368691],"study_design_scores_gemma":[0.0003100165,0.006481069,0.9199001,0.0003063799,0.00047698844,0.000534523,0.0052967137,0.042385887,0.010539828,0.005142826,0.0084796455,0.0001460115],"about_ca_topic_score_codex":0.0068112407,"about_ca_topic_score_gemma":0.010548627,"teacher_disagreement_score":0.9888774,"about_ca_system_score_codex":0.0017046958,"about_ca_system_score_gemma":0.0016146009,"threshold_uncertainty_score":0.05882263},"labels":[],"label_agreement":null},{"id":"W1550680615","doi":"10.1109/wpc.2005.31","title":"Software Clustering based on Omnipresent Object Detection","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Cluster analysis; Software; Object (grammar); Process (computing); Object detection; Data mining; Set (abstract data type); Artificial intelligence; Machine learning; Programming language","score_opus":0.016243895052569717,"score_gpt":0.2574559491668041,"score_spread":0.24121205411423438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1550680615","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05787505,0.00023798812,0.93903023,0.00010569001,0.000027461196,0.00011481748,0.00003383517,0.0012218177,0.0013530237],"genre_scores_gemma":[0.28859812,0.00022844304,0.70910645,0.00006946764,0.000033764445,0.000091140966,0.00017071355,0.00013086542,0.0015709843],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99771583,0.00036487993,0.00011286467,0.00052867713,0.0010905319,0.00018721804],"domain_scores_gemma":[0.9927078,0.0024476307,0.0010231392,0.0011602833,0.0024216194,0.00023949787],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016905239,0.00074829854,0.0010060458,0.0043244497,0.0011372713,0.0014380402,0.0020377133,0.0011194155,0.00064361264],"category_scores_gemma":[0.007937371,0.0005745553,0.00067260634,0.0019890866,0.0014080147,0.0020339168,0.0016526278,0.0008602037,0.0005116433],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003705165,0.00030553638,0.018547123,0.0003869081,0.00018682034,0.0005348485,0.0011264539,0.08398457,0.09438762,0.019102529,0.0034245143,0.77764255],"study_design_scores_gemma":[0.000035944227,0.00021733095,0.0091066435,0.00003434424,0.00009044395,0.0016850248,0.00024894954,0.8933215,0.07561113,0.013050022,0.0064416034,0.00015707282],"about_ca_topic_score_codex":0.0026687689,"about_ca_topic_score_gemma":0.004143982,"teacher_disagreement_score":0.0043244497,"about_ca_system_score_codex":0.0007455173,"about_ca_system_score_gemma":0.0010352969,"threshold_uncertainty_score":0.008940458},"labels":[],"label_agreement":null},{"id":"W1552527312","doi":"10.1007/978-0-85729-277-3_1","title":"The Role of Specification","year":2011,"lang":"en","type":"book-chapter","venue":"Texts in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software engineering; Software development; Domain (mathematical analysis); Software system; Business requirements; Documentation; Software; Engineering; Business process","score_opus":0.027187381759078505,"score_gpt":0.2506117309018594,"score_spread":0.2234243491427809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1552527312","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024370837,0.01779435,0.10423826,0.012929789,0.0011513485,0.00004141359,0.000116276744,0.0002716466,0.86101985],"genre_scores_gemma":[0.42152515,0.018599413,0.0707608,0.0074122306,0.0023699363,0.0004894457,0.0004514092,0.0010358404,0.4773558],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9976216,0.0013572819,0.0000915978,0.0003389476,0.0004386697,0.00015184862],"domain_scores_gemma":[0.9965963,0.0022860975,0.000117770556,0.0006200065,0.00027121126,0.000108690125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029435612,0.0011072451,0.0008235124,0.0013042103,0.0022148148,0.0075435066,0.001889134,0.002085489,0.01653655],"category_scores_gemma":[0.004660535,0.0008354353,0.0007119719,0.0016094276,0.018700631,0.011613076,0.0021010812,0.0066913045,0.0051021636],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000017033981,0.000002080896,0.0000067051433,0.000011849743,6.230703e-7,0.0000040789823,0.00015305537,0.00005332814,0.000020200334,0.9949993,0.0021818422,0.0025652924],"study_design_scores_gemma":[0.000004746355,0.000004491214,0.000021915384,0.000045654066,0.0000022025101,0.000028196697,0.00009631813,0.00036911003,0.000121651676,0.8945934,0.10470841,0.0000039055485],"about_ca_topic_score_codex":0.0021344998,"about_ca_topic_score_gemma":0.0016269116,"teacher_disagreement_score":0.01653655,"about_ca_system_score_codex":0.0039198175,"about_ca_system_score_gemma":0.0022393526,"threshold_uncertainty_score":0.055320323},"labels":[],"label_agreement":null},{"id":"W1552545845","doi":"10.1023/a:1011210102797","title":"Computing Cyclomatic Complexity with Cubic Flowgraphs","year":2001,"lang":"en","type":"article","venue":"Journal of Systems Integration","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Cyclomatic complexity; Rotation formalisms in three dimensions; Computer science; Decomposition; Computation; Theoretical computer science; Representation (politics); Component (thermodynamics); Software engineering; Software; Algorithm; Programming language; Mathematics","score_opus":0.032324135428177005,"score_gpt":0.27686078835088745,"score_spread":0.24453665292271046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1552545845","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48411578,0.0013229914,0.4484405,0.004714575,0.00031404372,0.00012875299,0.0023796523,0.0029869024,0.05559685],"genre_scores_gemma":[0.8894278,0.0006667699,0.09711039,0.00056674465,0.00037045032,0.0001247907,0.001673279,0.00052657083,0.009533165],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9991291,0.00019504246,0.00003741827,0.00020956955,0.00023290626,0.00019605902],"domain_scores_gemma":[0.9941912,0.0040534493,0.0003246285,0.0008508257,0.00029524037,0.00028465706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010741303,0.0008012475,0.0009511871,0.0017974456,0.001248607,0.0032103367,0.0017009921,0.0011125734,0.0137003185],"category_scores_gemma":[0.00896316,0.000659796,0.0010265078,0.00240113,0.0018196044,0.008352531,0.0020746684,0.0034513809,0.0011499642],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006030045,0.00019501761,0.003734873,0.00039056098,0.000094457726,0.00012264188,0.00033034672,0.13104632,0.0030921637,0.7605095,0.021675183,0.07820605],"study_design_scores_gemma":[0.00006371766,0.000033022414,0.0005037971,0.000022734881,0.00002630806,0.00004559145,0.00003886352,0.13892375,0.0011944579,0.8561126,0.0030165538,0.000018528577],"about_ca_topic_score_codex":0.0040338454,"about_ca_topic_score_gemma":0.006732431,"teacher_disagreement_score":0.0137003185,"about_ca_system_score_codex":0.0026940822,"about_ca_system_score_gemma":0.0012014145,"threshold_uncertainty_score":0.045832098},"labels":[],"label_agreement":null},{"id":"W1552571656","doi":"","title":"Software construction by composition of components","year":2010,"lang":"en","type":"article","venue":"Espace École de technologie supérieure (École de technologie supérieure)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Component-based software engineering; Software engineering; Software construction; Software development; Component (thermodynamics); Package development process; Software; Software framework; Software system; Programming language","score_opus":0.009675680740897419,"score_gpt":0.24735484883729772,"score_spread":0.2376791680964003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1552571656","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005583837,0.00037806932,0.9873766,0.00012701229,0.00007142418,0.00016018632,0.000020620388,0.001119832,0.0051624994],"genre_scores_gemma":[0.08001105,0.0008723363,0.9122156,0.000091018715,0.000051456926,0.00023265394,0.00024507183,0.00045312138,0.005827626],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99714345,0.00077503367,0.00022997361,0.0006246507,0.0010725792,0.000154232],"domain_scores_gemma":[0.997888,0.0006344809,0.00018140618,0.0007253515,0.00046195122,0.00010877663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020373526,0.0010930534,0.0007642679,0.0015758977,0.0012205708,0.0020299975,0.001632078,0.0009785687,0.0029924563],"category_scores_gemma":[0.0045085344,0.0008280235,0.0014970759,0.0013764559,0.001921793,0.0027269777,0.0040885746,0.001687799,0.0015766106],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019211692,0.00016935126,0.0028958055,0.0011315398,0.00021119746,0.0008842133,0.001839726,0.091195114,0.04915578,0.34514865,0.0048164637,0.50236],"study_design_scores_gemma":[0.00010196586,0.000366655,0.0010757634,0.00039532574,0.00038908783,0.0014040826,0.00039803516,0.39183268,0.046865955,0.29156372,0.2654935,0.00011326294],"about_ca_topic_score_codex":0.0014916721,"about_ca_topic_score_gemma":0.0011032178,"teacher_disagreement_score":0.0029924563,"about_ca_system_score_codex":0.0007862393,"about_ca_system_score_gemma":0.0021178236,"threshold_uncertainty_score":0.010774672},"labels":[],"label_agreement":null},{"id":"W1553310601","doi":"10.1109/wpc.2005.48","title":"Working Session on Interoperable Reengineering Services","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Interoperability; Business process reengineering; Semantic interoperability; Session (web analytics); Computer science; Cross-domain interoperability; Service (business); Software engineering; WS-I Basic Profile; Process (computing); Work (physics); Process management; World Wide Web; Web service; Engineering; Manufacturing engineering; Business","score_opus":0.016267076675126772,"score_gpt":0.2522153543497067,"score_spread":0.23594827767457993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1553310601","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019125009,0.048624244,0.40854317,0.091340765,0.036620975,0.0024583791,0.0036454168,0.004642586,0.38499957],"genre_scores_gemma":[0.06704186,0.05382071,0.15624247,0.011359598,0.01670457,0.0018913425,0.015724,0.0043483744,0.6728671],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99379,0.0015180886,0.0003540614,0.00079753704,0.0025006267,0.0010396633],"domain_scores_gemma":[0.99060506,0.0022738047,0.00029613933,0.0014803063,0.003207242,0.002137455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01791179,0.0019257437,0.0017939865,0.0024169497,0.0032353373,0.0079794675,0.0040698904,0.009078661,0.06523159],"category_scores_gemma":[0.009885463,0.0010170959,0.0023938902,0.0026282556,0.0014911572,0.012735202,0.007800968,0.0074295364,0.024903392],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038633376,0.0006680305,0.00082110096,0.0007910991,0.000099509925,0.0005278211,0.001791938,0.003809768,0.008098359,0.10503251,0.5551116,0.32286194],"study_design_scores_gemma":[0.000040935927,0.00014095534,0.00037265176,0.0004938739,0.000028219094,0.00019836548,0.00035863335,0.0021042994,0.0034295458,0.017766036,0.97502196,0.000044567947],"about_ca_topic_score_codex":0.004579671,"about_ca_topic_score_gemma":0.0042016893,"teacher_disagreement_score":0.06523159,"about_ca_system_score_codex":0.002050045,"about_ca_system_score_gemma":0.0062329657,"threshold_uncertainty_score":0.21822143},"labels":[],"label_agreement":null},{"id":"W1554594814","doi":"10.1109/metrics.2004.2","title":"A prototype empirical evaluation of test driven development","year":2004,"lang":"en","type":"article","venue":"IEEE International Software Metrics Symposium","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; University of Calgary","funders":"","keywords":"Computer science; Debugging; Test-driven development; Software development; Software development process; Software engineering; Test strategy; Software quality; New product development; Manual testing; Test (biology); Software; Empirical research; Software construction; Operating system; Business","score_opus":0.06676157119916065,"score_gpt":0.35087071105188666,"score_spread":0.284109139852726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1554594814","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88171625,0.0006869697,0.045176037,0.0023169888,0.00035353168,0.019247653,0.0012075808,0.0005235975,0.048771366],"genre_scores_gemma":[0.9193952,0.00035488405,0.060103107,0.00084804604,0.000091425645,0.015416953,0.00084107986,0.000094658615,0.0028546276],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9074475,0.07311607,0.003806659,0.0032193211,0.010746026,0.001664501],"domain_scores_gemma":[0.46356997,0.43139267,0.013585348,0.033247992,0.0525969,0.005607124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07918895,0.0007779818,0.00074038596,0.0021005385,0.0019224196,0.0032382093,0.0038315386,0.0025543885,0.008509725],"category_scores_gemma":[0.31063604,0.0006982367,0.0008333157,0.0024253677,0.0036788217,0.0050416696,0.003268574,0.0020953214,0.0016290792],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.023193065,0.081254244,0.1395811,0.008365877,0.0006916628,0.001765322,0.050385807,0.018637227,0.0146696605,0.06568946,0.027095517,0.5686711],"study_design_scores_gemma":[0.032729343,0.24066313,0.27981737,0.005055981,0.001227224,0.0019187469,0.0471492,0.09361948,0.026509412,0.048307132,0.22206,0.0009430012],"about_ca_topic_score_codex":0.002280847,"about_ca_topic_score_gemma":0.0018363554,"teacher_disagreement_score":0.07918895,"about_ca_system_score_codex":0.0043567894,"about_ca_system_score_gemma":0.004331052,"threshold_uncertainty_score":0.41879618},"labels":[],"label_agreement":null},{"id":"W1555024474","doi":"","title":"An experiment to investigate interacting versus nominal groups in software inspection","year":2003,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Fagan inspection; Software inspection; Computer science; Context (archaeology); Software; Psychology; Artificial intelligence; Software engineering; Software quality; Software development; Programming language","score_opus":0.13073143416064248,"score_gpt":0.4271523117467682,"score_spread":0.29642087758612573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1555024474","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9940969,0.000039934403,0.0015556948,0.00019161002,0.00007454812,0.00082471495,0.00004604749,0.000041765335,0.0031287905],"genre_scores_gemma":[0.9793188,0.00008193883,0.012849897,0.00046469434,0.00007352308,0.0037135365,0.0001324581,0.000028600374,0.0033364315],"study_design_codex":"nonrandomized_trial","study_design_gemma":"randomized_trial","domain_scores_codex":[0.99120104,0.0051522646,0.0007313028,0.001236951,0.0011009668,0.00057742314],"domain_scores_gemma":[0.91486466,0.07143455,0.0046452433,0.0040571187,0.0011267286,0.0038717343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011483754,0.00084068504,0.000785983,0.0005031058,0.0016942065,0.0018340634,0.0021422969,0.002893593,0.008280541],"category_scores_gemma":[0.03124692,0.0009075717,0.0005259705,0.00042855248,0.0022182253,0.0030049514,0.002073498,0.002699051,0.0009110316],"study_design_candidate":"randomized_trial","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.10384165,0.40756038,0.04862142,0.0020335878,0.00035056245,0.0017169833,0.101745985,0.0020312348,0.2155136,0.013549629,0.003379627,0.09965534],"study_design_scores_gemma":[0.039628103,0.72794664,0.10415848,0.0004183254,0.00050187897,0.0009453316,0.02937451,0.016291756,0.04400114,0.021678885,0.014547487,0.0005074408],"about_ca_topic_score_codex":0.00049294817,"about_ca_topic_score_gemma":0.0005854138,"teacher_disagreement_score":0.011483754,"about_ca_system_score_codex":0.0006658451,"about_ca_system_score_gemma":0.0010037777,"threshold_uncertainty_score":0.060732663},"labels":[],"label_agreement":null},{"id":"W1556141317","doi":"10.1007/978-3-642-37119-6_14","title":"RESource: A Framework for Online Matching of Assembly with Open Source Code","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; Concordia University","funders":"","keywords":"Computer science; Source code; Open source; Coding (social sciences); Software; Reverse engineering; File format; Set (abstract data type); Database; Programming language; Software engineering","score_opus":0.027590689778475667,"score_gpt":0.2978994570741178,"score_spread":0.2703087672956421,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1556141317","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010908343,0.00012392763,0.7810156,0.00012092393,0.0001431035,0.00014379957,0.0012587651,0.21181378,0.004289281],"genre_scores_gemma":[0.050104067,0.00035202387,0.8394475,0.00040567442,0.00018828182,0.0006863033,0.011899693,0.0819505,0.014965938],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99546635,0.00088343467,0.00047015745,0.00094706105,0.0017378374,0.0004951788],"domain_scores_gemma":[0.99179053,0.002892778,0.0004260935,0.003730628,0.00084950967,0.00031043545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004359914,0.0028048842,0.0017802345,0.004222728,0.0017489801,0.0052457545,0.008105242,0.0031050593,0.06912295],"category_scores_gemma":[0.020972744,0.0032174904,0.004617659,0.003473581,0.00176553,0.011229774,0.009885021,0.003992549,0.040734895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013275638,0.00041094286,0.0027487776,0.0013592574,0.00032369694,0.00061831955,0.0008927975,0.018515931,0.012839349,0.14055094,0.24416988,0.5762425],"study_design_scores_gemma":[0.00033354398,0.00017260226,0.0011052215,0.00040424388,0.0001573145,0.00070310227,0.0003200398,0.3651965,0.057305068,0.24356604,0.3304322,0.00030412117],"about_ca_topic_score_codex":0.005322657,"about_ca_topic_score_gemma":0.005315696,"teacher_disagreement_score":0.06912295,"about_ca_system_score_codex":0.0016150922,"about_ca_system_score_gemma":0.002701594,"threshold_uncertainty_score":0.23123926},"labels":[],"label_agreement":null},{"id":"W1556736612","doi":"10.1007/978-3-540-30203-2_1","title":"Improving Flow in Software Development Through Graphical Representations","year":2004,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software development; Software engineering; Software; Software construction; Program comprehension; Software visualization; Software analytics; Software framework; Human–computer interaction; Software system; Programming language","score_opus":0.01993614611788694,"score_gpt":0.26879061743768595,"score_spread":0.24885447131979901,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1556736612","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0056460905,0.00064683007,0.9806778,0.0002853642,0.0000636651,0.000033234843,0.000031583306,0.0039822333,0.008633243],"genre_scores_gemma":[0.12097891,0.0022968322,0.85700035,0.000147352,0.00009754598,0.00013066472,0.00026824317,0.001583209,0.017496927],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99915004,0.00028948358,0.000056263154,0.00011864485,0.00031777273,0.00006788405],"domain_scores_gemma":[0.99751616,0.0015512672,0.00014783807,0.0005051548,0.00021910667,0.00006045666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001110451,0.0010249849,0.00044614667,0.0013044359,0.00034942708,0.0017716476,0.0012119054,0.00091586506,0.007582457],"category_scores_gemma":[0.0060704984,0.00066406716,0.0007834365,0.0013186642,0.0008749453,0.004317401,0.0014278737,0.0019626485,0.0018557751],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011600769,0.00012870354,0.00040116045,0.00043941598,0.000027861395,0.00007885239,0.00063768774,0.04417987,0.016154883,0.13947333,0.008615199,0.78974706],"study_design_scores_gemma":[0.0001474893,0.00035452362,0.00070273515,0.00064326014,0.00018937007,0.00040837165,0.0001785603,0.40252528,0.07013407,0.38285866,0.141759,0.00009871423],"about_ca_topic_score_codex":0.00080389326,"about_ca_topic_score_gemma":0.0006735065,"teacher_disagreement_score":0.007582457,"about_ca_system_score_codex":0.00046220733,"about_ca_system_score_gemma":0.0006090945,"threshold_uncertainty_score":0.02536583},"labels":[],"label_agreement":null},{"id":"W1557076420","doi":"","title":"An empirical evaluation of system and regression testing","year":2002,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Regression testing; Computer science; Software quality; Empirical research; Context (archaeology); Non-regression testing; Quality (philosophy); Software reliability testing; Software engineering; Development testing; Reliability engineering; Software; Software system; Software development; Software construction; Engineering; Statistics; Programming language; Mathematics","score_opus":0.3136702464333212,"score_gpt":0.47068405716956496,"score_spread":0.15701381073624376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1557076420","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9561801,0.0020647596,0.020181058,0.00150493,0.00012852138,0.0008422612,0.0009391722,0.00026324493,0.017895872],"genre_scores_gemma":[0.99082315,0.000241427,0.0068996237,0.00020213847,0.00009566137,0.0004337757,0.00059959973,0.00005610726,0.00064839673],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.80564255,0.16187894,0.0070289895,0.0048355516,0.019109027,0.0015049946],"domain_scores_gemma":[0.15141383,0.7671471,0.032460008,0.022253653,0.024006896,0.0027184382],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10056284,0.00079426816,0.0006634504,0.004562589,0.0010117211,0.00259981,0.0022657595,0.0019371932,0.004837502],"category_scores_gemma":[0.5158507,0.00040978208,0.000699081,0.0039764754,0.0038583516,0.0058441786,0.002404916,0.0018673411,0.0010322867],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008122681,0.006736598,0.74346304,0.0013093724,0.0010851134,0.00060441234,0.005080531,0.017016105,0.0012209895,0.01704777,0.0057582995,0.19255508],"study_design_scores_gemma":[0.0020009032,0.029658338,0.7206479,0.00092131866,0.0008809093,0.002702165,0.0070507196,0.19755027,0.0047414545,0.011909001,0.02167602,0.0002609487],"about_ca_topic_score_codex":0.0022632806,"about_ca_topic_score_gemma":0.0018152265,"teacher_disagreement_score":0.10056284,"about_ca_system_score_codex":0.002183749,"about_ca_system_score_gemma":0.0012279991,"threshold_uncertainty_score":0.5318335},"labels":[],"label_agreement":null},{"id":"W1562806970","doi":"10.1007/978-3-642-05441-9_20","title":"Evaluation of Test-Driven Development: An Academic Case Study","year":2009,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Algoma University","funders":"","keywords":"Test (biology); Cohesion (chemistry); Mathematics education; Group (periodic table); Computer science; Test-driven development; Simple (philosophy); Quality (philosophy); Group cohesiveness; Code (set theory); Engineering management; Software engineering; Psychology; Engineering; Programming language; Social psychology; Software development","score_opus":0.2791864633978125,"score_gpt":0.45500680195127535,"score_spread":0.17582033855346285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1562806970","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8896373,0.0010564592,0.07295612,0.0011512173,0.00006903642,0.0011222165,0.0004273955,0.00088224554,0.03269805],"genre_scores_gemma":[0.95548564,0.000314723,0.03911619,0.00012256303,0.000015544589,0.0003028443,0.00045965894,0.00013584772,0.004046984],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9899339,0.006171267,0.00044688454,0.00038094545,0.0025896693,0.00047719173],"domain_scores_gemma":[0.95305544,0.03300746,0.001517772,0.003770728,0.0069087674,0.0017397758],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010127596,0.0007142731,0.00045194497,0.0014042939,0.0010199526,0.001904864,0.002602903,0.0017034035,0.0018128115],"category_scores_gemma":[0.03469262,0.00029003914,0.0004135151,0.0015635957,0.001101641,0.0011407927,0.0015204231,0.0009262515,0.00060584646],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032191651,0.011144034,0.033208847,0.0017276757,0.00020708071,0.006768738,0.01732835,0.07188315,0.024102101,0.016419291,0.015620991,0.7983705],"study_design_scores_gemma":[0.0024461383,0.02275775,0.07916896,0.001665478,0.00055495364,0.010530834,0.021360656,0.5558974,0.14992157,0.02739757,0.1278296,0.00046903593],"about_ca_topic_score_codex":0.0043507167,"about_ca_topic_score_gemma":0.005671516,"teacher_disagreement_score":0.010127596,"about_ca_system_score_codex":0.0020471702,"about_ca_system_score_gemma":0.0019456272,"threshold_uncertainty_score":0.053560436},"labels":[],"label_agreement":null},{"id":"W1563960834","doi":"","title":"A case study in the use of defect classification in inspections","year":2001,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Classification scheme; Computer science; Software bug; Software inspection; Software metric; Metric (unit); Scheme (mathematics); Software; Data mining; Variety (cybernetics); Machine learning; Reliability engineering; Software development; Software engineering; Software quality; Artificial intelligence; Engineering; Mathematics; Operations management","score_opus":0.3660798221192619,"score_gpt":0.45170654738677996,"score_spread":0.08562672526751808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1563960834","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89678216,0.0016842845,0.07106623,0.0033385286,0.00015218041,0.0014091223,0.00023107091,0.00019226824,0.025144093],"genre_scores_gemma":[0.8819291,0.00084428163,0.11322019,0.00042075277,0.000044536042,0.00046981624,0.00010547106,0.000056057102,0.0029096622],"study_design_codex":"design_other","study_design_gemma":"case_report","domain_scores_codex":[0.9605932,0.030764507,0.0014478415,0.0013198352,0.0051433737,0.00073115394],"domain_scores_gemma":[0.903234,0.080256455,0.0036216856,0.0048610037,0.0065878313,0.0014389841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019561147,0.00055854645,0.00061720976,0.0036730743,0.004400073,0.0029676557,0.0022479617,0.0047139423,0.0011357382],"category_scores_gemma":[0.050802864,0.0005624078,0.00074906653,0.0038409359,0.003427846,0.0028828185,0.0030436823,0.0021786513,0.00035305283],"study_design_candidate":"case_report","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002176402,0.008579132,0.20543315,0.0026321611,0.00023166716,0.07343323,0.22811234,0.02597247,0.024105685,0.055776462,0.010772949,0.36277437],"study_design_scores_gemma":[0.0010425014,0.014111432,0.23425077,0.0042717922,0.0005745513,0.089056924,0.23881528,0.14730176,0.059922185,0.033783138,0.1758707,0.0009989729],"about_ca_topic_score_codex":0.008879937,"about_ca_topic_score_gemma":0.018723957,"teacher_disagreement_score":0.019561147,"about_ca_system_score_codex":0.0034112355,"about_ca_system_score_gemma":0.0014920819,"threshold_uncertainty_score":0.10345042},"labels":[],"label_agreement":null},{"id":"W1564596608","doi":"10.1016/s0065-2458(00)80008-x","title":"An empirical review of software process assessments","year":2000,"lang":"en","type":"book-chapter","venue":"Advances in computers","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Process (computing); Computer science; Software; Work (physics); Process management; Risk analysis (engineering); Software engineering; Management science; Engineering; Business","score_opus":0.026940248686243985,"score_gpt":0.38518763211246326,"score_spread":0.35824738342621926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1564596608","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010974909,0.9543091,0.0071708625,0.005586383,0.00049931876,0.000071570605,0.0007131875,0.00006829713,0.020606356],"genre_scores_gemma":[0.11297241,0.86899865,0.009451873,0.0020277207,0.0006531813,0.00013350966,0.00113107,0.000100405974,0.0045312406],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.98428935,0.00601706,0.0016917327,0.0006795074,0.0071078953,0.00021440395],"domain_scores_gemma":[0.8079256,0.15636986,0.006688142,0.0025117649,0.02600708,0.0004975858],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014563549,0.0004314982,0.0007589886,0.013010983,0.00062577805,0.0030270088,0.0011095866,0.00077673956,0.004199501],"category_scores_gemma":[0.11808115,0.00041126058,0.00046204345,0.022056164,0.0013590092,0.005572366,0.001072056,0.0012836196,0.0013346729],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010714356,0.000108933215,0.007812134,0.009578862,0.00009423555,0.00007270793,0.0007368169,0.00038671185,0.00041505025,0.017638795,0.03911007,0.9239385],"study_design_scores_gemma":[0.00006055048,0.00047598802,0.114884794,0.05222166,0.0011469054,0.0012260244,0.0042334227,0.0020485595,0.0035192913,0.030946193,0.78912723,0.000109443004],"about_ca_topic_score_codex":0.005460318,"about_ca_topic_score_gemma":0.0096107,"teacher_disagreement_score":0.98543644,"about_ca_system_score_codex":0.002839835,"about_ca_system_score_gemma":0.006828248,"threshold_uncertainty_score":0.07702029},"labels":[],"label_agreement":null},{"id":"W1567965530","doi":"10.1007/978-3-540-85502-6_20","title":"Cases, Predictions, and Accuracy Learning and Its Application to Effort Estimation","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Analogy; Estimation; Set (abstract data type); Similarity (geometry); Quality (philosophy); Reliability (semiconductor); Machine learning; Process (computing); Artificial intelligence; Data mining; Software quality; Software; Software development","score_opus":0.01468652076736102,"score_gpt":0.27523563037040866,"score_spread":0.2605491096030476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1567965530","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043543085,0.0020139604,0.9460202,0.0014414816,0.00013475025,0.000083897736,0.00023118517,0.0006329652,0.005898464],"genre_scores_gemma":[0.7264241,0.001606666,0.26382414,0.00023628525,0.00051555084,0.00029023088,0.00060080027,0.00016893697,0.006333421],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952251,0.0025576046,0.00027609803,0.00076373015,0.00096421735,0.00021325072],"domain_scores_gemma":[0.91691786,0.074833915,0.0018535048,0.0032685758,0.0026134443,0.0005126819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010652044,0.0012701935,0.0017635091,0.0036489847,0.0009685046,0.0029869215,0.0033587518,0.0020605323,0.004435982],"category_scores_gemma":[0.087453686,0.0008861041,0.0013741683,0.0036814304,0.00301676,0.007624366,0.0026497797,0.0034973158,0.0005621189],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004461357,0.00033718863,0.023259655,0.00026074014,0.00019849766,0.00017787988,0.0007822373,0.25710607,0.00051562383,0.23796296,0.0092681395,0.46968496],"study_design_scores_gemma":[0.000013664372,0.000037370763,0.0016110879,0.000029091421,0.000033426208,0.00006160257,0.0000625144,0.76048225,0.000337177,0.23644586,0.00086222123,0.000023733855],"about_ca_topic_score_codex":0.007545709,"about_ca_topic_score_gemma":0.0061674546,"teacher_disagreement_score":0.010652044,"about_ca_system_score_codex":0.001619181,"about_ca_system_score_gemma":0.0011605655,"threshold_uncertainty_score":0.05633408},"labels":[],"label_agreement":null},{"id":"W1569329269","doi":"10.1007/3-540-44704-0_17","title":"QF2D: A Different Way to Measure Software Quality","year":2001,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Quality function deployment; Computer science; Dimension (graph theory); Software quality; Context (archaeology); Quality (philosophy); Viewpoints; Benchmarking; Software engineering; Software; Software development; New product development; Mathematics; Programming language","score_opus":0.039982336861006204,"score_gpt":0.2895074407310006,"score_spread":0.24952510386999438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1569329269","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044227403,0.0012149842,0.9180805,0.001045111,0.0006512695,0.00023166763,0.0043781414,0.006690844,0.02347995],"genre_scores_gemma":[0.3855243,0.00058544363,0.5864264,0.0009111981,0.00021598284,0.0006415,0.0036923292,0.0016642767,0.020338576],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9904861,0.0016057706,0.00061131077,0.00096692884,0.0059268153,0.00040315912],"domain_scores_gemma":[0.9852937,0.0069633787,0.0011623248,0.0021147605,0.0040428806,0.000422967],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042021223,0.0011827535,0.00094401446,0.007986007,0.0005723136,0.0034615132,0.0014481554,0.0019329387,0.008181244],"category_scores_gemma":[0.024741758,0.00038367382,0.0010475856,0.005702408,0.0010042536,0.0052526486,0.0018140898,0.0014567717,0.002261894],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043571298,0.0002918608,0.049343567,0.0010205064,0.00037889514,0.00012673522,0.00096926786,0.00928229,0.015250974,0.050521765,0.04148763,0.83089083],"study_design_scores_gemma":[0.00035975312,0.0018691479,0.11909786,0.0007321214,0.00059155794,0.0023240978,0.0013414907,0.29150862,0.1347923,0.16131933,0.28514573,0.0009180473],"about_ca_topic_score_codex":0.003385189,"about_ca_topic_score_gemma":0.004076825,"teacher_disagreement_score":0.008181244,"about_ca_system_score_codex":0.0013354174,"about_ca_system_score_gemma":0.0007423472,"threshold_uncertainty_score":0.027368963},"labels":[],"label_agreement":null},{"id":"W1572112458","doi":"10.1109/apsec.2004.13","title":"A Systematic Study of UML Class Diagram Constituents for their Abstract and Precise Recovery","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Class diagram; Computer science; Programming language; Applications of UML; UML tool; Java; Reverse engineering; Unified Modeling Language; Suite; Class (philosophy); Communication diagram; Software engineering; Source code; Software; Artificial intelligence","score_opus":0.02401572227920324,"score_gpt":0.278695311654755,"score_spread":0.25467958937555174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1572112458","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023813592,0.0019074124,0.9697992,0.00035215728,0.000042885604,0.00043290007,0.00019765417,0.00084806635,0.0026060648],"genre_scores_gemma":[0.09450691,0.0021781519,0.8999569,0.00017140605,0.00003444776,0.00041620265,0.00082003395,0.0006023268,0.0013136112],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9807483,0.007311473,0.0022260987,0.0023826612,0.006994074,0.0003375087],"domain_scores_gemma":[0.9107019,0.04682355,0.0071751988,0.021403814,0.013403464,0.00049209537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011838129,0.00093847944,0.0010854101,0.006652125,0.0022376655,0.0037483752,0.0014613756,0.0012442801,0.0014496096],"category_scores_gemma":[0.08022043,0.0014107809,0.0016041022,0.005426807,0.0024611226,0.008978044,0.0017879274,0.0031768417,0.00074095244],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014002944,0.00039825658,0.0105343135,0.002433809,0.00010039563,0.0006805305,0.005285332,0.011140543,0.047562465,0.13178399,0.0035518876,0.7863884],"study_design_scores_gemma":[0.000106559404,0.0007011735,0.023781179,0.0047151013,0.00052586885,0.009641328,0.004059014,0.1543064,0.25104967,0.24251817,0.3080081,0.0005874281],"about_ca_topic_score_codex":0.0013614026,"about_ca_topic_score_gemma":0.0032044777,"teacher_disagreement_score":0.011838129,"about_ca_system_score_codex":0.0013172877,"about_ca_system_score_gemma":0.0045784013,"threshold_uncertainty_score":0.06260675},"labels":[],"label_agreement":null},{"id":"W1574403820","doi":"10.1007/3-540-45341-5_30","title":"Design and Implementation of a UML-Based Design Repository","year":2001,"lang":"en","type":"book-chapter","venue":"Notes on numerical fluid mechanics and multidisciplinary design","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Unified Modeling Language; Metamodeling; Computer science; Software engineering; Applications of UML; UML tool; Systems engineering; Software; Engineering; Programming language","score_opus":0.03907834946473531,"score_gpt":0.29198684676792386,"score_spread":0.25290849730318854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1574403820","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049407138,0.00009107462,0.95771724,0.00031546352,0.00009523098,0.0007400724,0.00045150024,0.032986112,0.0026626019],"genre_scores_gemma":[0.04315619,0.00024485757,0.9390053,0.000341264,0.000038358496,0.0010685957,0.0031403461,0.006587126,0.0064179054],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99168384,0.0020721168,0.0014057717,0.0011881703,0.0031233116,0.00052686356],"domain_scores_gemma":[0.9863137,0.004610066,0.00092653325,0.0045336187,0.002770677,0.0008453972],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017646117,0.0016770118,0.0017517153,0.0040179933,0.0012517364,0.009669983,0.008692451,0.0033341823,0.014343096],"category_scores_gemma":[0.025698308,0.00322755,0.0025239775,0.0022070506,0.0017352963,0.007476309,0.0050951247,0.004753449,0.0071455254],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014108212,0.0020319023,0.0057046353,0.0022486025,0.00059062627,0.0016997866,0.0062840944,0.059184086,0.06794312,0.106843784,0.04439362,0.70166487],"study_design_scores_gemma":[0.0015969246,0.0008320228,0.0021450364,0.0011374941,0.00079741376,0.0014913421,0.0010374598,0.52636945,0.14493562,0.037547763,0.28164798,0.00046143765],"about_ca_topic_score_codex":0.0035092519,"about_ca_topic_score_gemma":0.002759035,"teacher_disagreement_score":0.017646117,"about_ca_system_score_codex":0.0016871699,"about_ca_system_score_gemma":0.005297812,"threshold_uncertainty_score":0.093322694},"labels":[],"label_agreement":null},{"id":"W1576661435","doi":"","title":"Data quality by design: a goal-oriented approach","year":2010,"lang":"en","type":"dissertation","venue":"TSpace (University of Toronto)","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Conceptual schema; Quality assurance; Software engineering; Data quality; Data mining; Database; Engineering; Metric (unit)","score_opus":0.03975536040037669,"score_gpt":0.3050509463482356,"score_spread":0.2652955859478589,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1576661435","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017696735,0.0022740758,0.9608497,0.007853441,0.00013141424,0.00038986673,0.000072771145,0.00022436351,0.02643472],"genre_scores_gemma":[0.06364502,0.0035130403,0.9209554,0.0020510983,0.00019656321,0.0012004741,0.0002330287,0.00024668095,0.0079586925],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97880906,0.01260952,0.0014584891,0.0018510497,0.004464081,0.0008077488],"domain_scores_gemma":[0.97929937,0.012746241,0.0011942248,0.002813507,0.0031096253,0.00083697663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03180521,0.0025939913,0.0015981125,0.0065155416,0.0027863123,0.018061254,0.0068899384,0.004927752,0.0037365633],"category_scores_gemma":[0.018929722,0.001885734,0.0024429243,0.0038078446,0.018609485,0.011251524,0.0064300275,0.0061048428,0.0014006364],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00001224538,0.00006589719,0.00032554922,0.00042173197,0.00004946691,0.000109174565,0.0016252395,0.006139237,0.0004103408,0.962667,0.0017790062,0.026395133],"study_design_scores_gemma":[0.000040138526,0.00008171495,0.000220513,0.00071788597,0.00007284351,0.00018872955,0.0015479836,0.02110623,0.0014480577,0.87599915,0.09852109,0.000055706736],"about_ca_topic_score_codex":0.0049017156,"about_ca_topic_score_gemma":0.0035617456,"teacher_disagreement_score":0.03180521,"about_ca_system_score_codex":0.009591987,"about_ca_system_score_gemma":0.012379995,"threshold_uncertainty_score":0.16820401},"labels":[],"label_agreement":null},{"id":"W1576705238","doi":"","title":"Visualization Techniques for Program ComprehensionA Literature Review","year":2006,"lang":"en","type":"article","venue":"New Trends in Software Methodologies, Tools and Techniques","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada","funders":"","keywords":"Visualization; Computer science; Software visualization; Field (mathematics); Software engineering; Component (thermodynamics); Point (geometry); Software; Data science; Key (lock); Software development; Information visualization; Data visualization; Component-based software engineering; Data mining; Programming language","score_opus":0.11033280646145692,"score_gpt":0.4198981169595048,"score_spread":0.30956531049804786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1576705238","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022061619,0.975747,0.012343551,0.0011495063,0.0003019572,0.000059974813,0.00017410317,0.00032437447,0.0076933615],"genre_scores_gemma":[0.015221146,0.9598662,0.020507047,0.00030738686,0.00035718453,0.00012309912,0.00044269953,0.00013329127,0.0030419158],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972059,0.00087006163,0.0003002915,0.0003598853,0.0011283577,0.00013546394],"domain_scores_gemma":[0.9875723,0.007914924,0.00076397764,0.00036519478,0.0032321562,0.00015153889],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032954924,0.0014201872,0.0015933439,0.011403832,0.0009308507,0.0034335912,0.0014486678,0.001261646,0.009280343],"category_scores_gemma":[0.0164629,0.0007080316,0.0013245814,0.015559336,0.0007522868,0.005433474,0.0010135816,0.00109695,0.002318415],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000056935616,0.00008481069,0.0007116647,0.014836937,0.00010546308,0.00013634666,0.00053761195,0.0007693062,0.0012103533,0.0046732905,0.016626725,0.9602506],"study_design_scores_gemma":[0.00006359586,0.0003436983,0.01124064,0.050643858,0.0008532801,0.0019806314,0.0022145184,0.005782473,0.0073228343,0.024281137,0.89509904,0.00017423197],"about_ca_topic_score_codex":0.0046800897,"about_ca_topic_score_gemma":0.00518974,"teacher_disagreement_score":0.011403832,"about_ca_system_score_codex":0.0012115354,"about_ca_system_score_gemma":0.0030426204,"threshold_uncertainty_score":0.031045854},"labels":[],"label_agreement":null},{"id":"W1576984811","doi":"10.1023/a:1025320418915","title":"An Externally Replicated Experiment for Evaluating the Learning Effectiveness of Using Simulations in Software Project Management Education","year":2003,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"European Commission; Technische Universität Kaiserslautern","keywords":"Computer science; COCOMO; Debriefing; Software project management; Empirical research; Project management; Software; Software development; Process (computing); Software engineering; Knowledge management; Engineering management; Engineering; Systems engineering; Software construction; Psychology","score_opus":0.07108431282324686,"score_gpt":0.42070935125680736,"score_spread":0.3496250384335605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1576984811","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98685294,0.000028791228,0.0077241347,0.0000708632,0.00016370654,0.0027957668,0.00015744075,0.000146323,0.0020599824],"genre_scores_gemma":[0.96532184,0.00005926689,0.02300745,0.00013238749,0.00008900061,0.008364758,0.00024759932,0.000063331194,0.0027145648],"study_design_codex":"nonrandomized_trial","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9845834,0.0095525,0.0015241839,0.002010467,0.001831183,0.00049822475],"domain_scores_gemma":[0.8609152,0.099177174,0.008730609,0.019781072,0.007433842,0.0039620693],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015933555,0.001635989,0.0013257804,0.00077476434,0.001271505,0.0016395631,0.002849667,0.0031940935,0.004298294],"category_scores_gemma":[0.0868552,0.0009666419,0.0007350606,0.00052704645,0.0017592416,0.0017354682,0.0018425207,0.0027508459,0.00097351545],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.1841163,0.46683684,0.027218943,0.0013742907,0.00078179216,0.0003896489,0.006694447,0.026258517,0.17149977,0.0048180963,0.0015944771,0.108416796],"study_design_scores_gemma":[0.06618019,0.69301736,0.049552824,0.00026599452,0.0011623172,0.000236074,0.0013713664,0.042075746,0.13439745,0.0057447013,0.0055871923,0.0004087381],"about_ca_topic_score_codex":0.0007517543,"about_ca_topic_score_gemma":0.000722752,"teacher_disagreement_score":0.9840664,"about_ca_system_score_codex":0.0010010907,"about_ca_system_score_gemma":0.0023216235,"threshold_uncertainty_score":0.08426571},"labels":[],"label_agreement":null},{"id":"W1579515608","doi":"10.1016/s0065-2458(10)79005-7","title":"The Tools Perspective on Software Reverse Engineering","year":2010,"lang":"en","type":"book-chapter","venue":"Advances in computers","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Reverse engineering; Computer science; Software engineering; Software; Abstraction; Systems engineering; Engineering","score_opus":0.01267546753794807,"score_gpt":0.26424077124964757,"score_spread":0.2515653037116995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1579515608","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019189002,0.037665673,0.19796702,0.0066064848,0.0026667423,0.000037690104,0.00010164602,0.00063739286,0.7523985],"genre_scores_gemma":[0.08075742,0.07645825,0.15397647,0.003902299,0.002778364,0.00018473042,0.00034934026,0.0011342505,0.68045896],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.999257,0.00022031977,0.00003601931,0.00010104295,0.00033385743,0.00005173645],"domain_scores_gemma":[0.9984602,0.0010474379,0.000050212042,0.000222962,0.00017768708,0.000041493036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00077744713,0.0010414093,0.0005361655,0.0018832952,0.0009242799,0.005042887,0.0013875344,0.0017093216,0.014119221],"category_scores_gemma":[0.0019012693,0.00057411025,0.00066534494,0.0023595553,0.0042853644,0.007903481,0.0014148187,0.004724572,0.0065111937],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000039113656,0.0000087160615,0.000018344153,0.000082530314,0.0000031602,0.000055317527,0.00021493572,0.0004519557,0.00025256252,0.94013894,0.016279144,0.042490292],"study_design_scores_gemma":[0.000003961196,0.000011501234,0.000029878884,0.00018281421,0.000005530528,0.00023618004,0.00011588074,0.0009873159,0.0006457413,0.540717,0.45705587,0.000008349519],"about_ca_topic_score_codex":0.0013119533,"about_ca_topic_score_gemma":0.002019086,"teacher_disagreement_score":0.014119221,"about_ca_system_score_codex":0.0014858362,"about_ca_system_score_gemma":0.0015126516,"threshold_uncertainty_score":0.047233462},"labels":[],"label_agreement":null},{"id":"W1580134275","doi":"10.1007/978-3-540-78743-3_20","title":"A Domain Analysis to Specify Design Defects and Generate Detection Algorithms","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Domain (mathematical analysis); Template; Process (computing); Algorithm; Precision and recall; Domain analysis; Quality (philosophy); Software; Programming language; Data mining; Artificial intelligence; Software system; Software construction","score_opus":0.02386372752955233,"score_gpt":0.24852376944357105,"score_spread":0.22466004191401873,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1580134275","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010695393,0.00003229053,0.9958752,0.0000448493,0.000011407808,0.00007501281,0.00013299388,0.0016437588,0.0011149357],"genre_scores_gemma":[0.035955474,0.0001897031,0.95941967,0.00011638063,0.000027329532,0.00023303057,0.0008910474,0.0008303021,0.0023370893],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998016,0.0004780349,0.00020512159,0.00031348466,0.00082867884,0.00015867931],"domain_scores_gemma":[0.9945437,0.0030483077,0.00026489867,0.0009675277,0.0011100994,0.00006551732],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020005552,0.0022325772,0.0011451402,0.0031300026,0.0009768179,0.002389151,0.0018753409,0.0017048424,0.0069089727],"category_scores_gemma":[0.007890726,0.0013620328,0.0038625617,0.0014061307,0.0011112817,0.002676737,0.0017431041,0.003106007,0.0037485694],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017074974,0.00041264063,0.0027416763,0.0011440447,0.00019854851,0.0008204226,0.00056528347,0.23859233,0.040127523,0.20954928,0.019917423,0.48576],"study_design_scores_gemma":[0.00004482862,0.000095619514,0.00037465928,0.00024125601,0.00013470014,0.00057508436,0.00015439384,0.81350696,0.0329066,0.12647463,0.025437897,0.000053350464],"about_ca_topic_score_codex":0.0022226782,"about_ca_topic_score_gemma":0.002777432,"teacher_disagreement_score":0.0069089727,"about_ca_system_score_codex":0.0009467436,"about_ca_system_score_gemma":0.0019620836,"threshold_uncertainty_score":0.023112833},"labels":[],"label_agreement":null},{"id":"W1581639473","doi":"10.1007/s10664-012-9209-9","title":"Automated topic naming","year":2012,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of British Columbia; University of Alberta","funders":"","keywords":"Computer science; Latent Dirichlet allocation; Topic model; Commit; Categorization; Software; Context (archaeology); Information retrieval; Domain (mathematical analysis); Artificial intelligence; Natural language processing; Data science; Database; Programming language","score_opus":0.0256654668367926,"score_gpt":0.2953701151461752,"score_spread":0.2697046483093826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1581639473","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02163064,0.0011782096,0.8201336,0.0010909581,0.0013558154,0.0008249888,0.015188715,0.111389354,0.02720772],"genre_scores_gemma":[0.16100122,0.0006583045,0.7594329,0.00033076367,0.0006044373,0.00074363506,0.039361235,0.009452349,0.028415095],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993664,0.0017912894,0.0007036392,0.0017550412,0.0015183033,0.00056766754],"domain_scores_gemma":[0.98501855,0.0051109344,0.0006881762,0.0046881256,0.0037404753,0.0007536874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003726395,0.0017770369,0.002168254,0.0102029415,0.0035137648,0.0067146043,0.0024372602,0.0017481969,0.05306211],"category_scores_gemma":[0.019340817,0.0011083531,0.0025482895,0.0059882747,0.0008304082,0.008740413,0.006092536,0.0022251224,0.036676116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006562417,0.00018027984,0.0041445782,0.0007872213,0.0001544787,0.00026513234,0.0013302565,0.001746112,0.038478833,0.032340344,0.13964218,0.7802744],"study_design_scores_gemma":[0.00027359443,0.00020470943,0.00797001,0.00030973024,0.00037075125,0.0015230809,0.0032590395,0.19447966,0.091416515,0.13495828,0.56495285,0.0002819064],"about_ca_topic_score_codex":0.0032494674,"about_ca_topic_score_gemma":0.0050712693,"teacher_disagreement_score":0.05306211,"about_ca_system_score_codex":0.0017552067,"about_ca_system_score_gemma":0.0035841744,"threshold_uncertainty_score":0.17751044},"labels":[],"label_agreement":null},{"id":"W1591093410","doi":"","title":"Efficient mapping of software system traces to architectural views","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; TRACE (psycholinguistics); Viewpoints; Software visualization; Encoding (memory); Software engineering; Software; Variety (cybernetics); Software system; Software development; Software construction; Programming language; Artificial intelligence","score_opus":0.021718302857512597,"score_gpt":0.25303717563248923,"score_spread":0.23131887277497665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1591093410","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032032896,0.00016313297,0.9526714,0.00031040533,0.00003321791,0.00012537611,0.0008745601,0.012048298,0.0017407323],"genre_scores_gemma":[0.30680043,0.0006036254,0.6845464,0.00006377534,0.000031304608,0.0002748763,0.003001633,0.0025935175,0.0020844152],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835324,0.00045704865,0.00015758089,0.00023096608,0.0006626701,0.0001385813],"domain_scores_gemma":[0.99147415,0.004224989,0.0008144951,0.002256636,0.0010461002,0.00018366968],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015185905,0.0010592797,0.0008016641,0.002838655,0.0006514694,0.003206137,0.001509443,0.0010137303,0.0035156002],"category_scores_gemma":[0.019625045,0.0010441148,0.00071541074,0.003091718,0.00073114794,0.004956151,0.0027652977,0.0019417743,0.0009025314],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077686436,0.00023679403,0.0071421945,0.0007347608,0.00007207327,0.0007335373,0.00410193,0.100131415,0.04282337,0.13852732,0.010344632,0.6943751],"study_design_scores_gemma":[0.00009431474,0.0001970553,0.0024772673,0.0001672795,0.000059990187,0.00038880814,0.0009605101,0.74685144,0.078881,0.13725945,0.032546222,0.000116612646],"about_ca_topic_score_codex":0.0052603832,"about_ca_topic_score_gemma":0.004611471,"teacher_disagreement_score":0.0052603832,"about_ca_system_score_codex":0.0012452785,"about_ca_system_score_gemma":0.0017287171,"threshold_uncertainty_score":0.011760831},"labels":[],"label_agreement":null},{"id":"W1591116488","doi":"10.21236/ada390602","title":"On the Role of Randomization in Software Engineering","year":2001,"lang":"en","type":"report","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Office of Naval Research; Natural Sciences and Engineering Research Council of Canada","keywords":"Randomization; Software; Computer science; Software engineering; Programming language; Biology; Bioinformatics; Clinical trial","score_opus":0.015132021588664614,"score_gpt":0.24979260712021295,"score_spread":0.23466058553154834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1591116488","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024154767,0.00923465,0.8806717,0.022514116,0.00048365298,0.00010588355,0.000104724844,0.0009126754,0.061817802],"genre_scores_gemma":[0.6704389,0.014590925,0.29158232,0.0046871533,0.0018008598,0.000487236,0.00019808204,0.0005488077,0.01566565],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98744875,0.0068268063,0.00070854754,0.0014649775,0.0030551057,0.00049584295],"domain_scores_gemma":[0.94298345,0.046641324,0.0021152443,0.0058669383,0.0018060976,0.00058686506],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012106651,0.00074319803,0.0010099323,0.0020997939,0.002056516,0.004227601,0.0011228187,0.003139451,0.0032039315],"category_scores_gemma":[0.042037513,0.0006247376,0.0007492865,0.0028487255,0.020792283,0.014209334,0.003931292,0.00415204,0.0011063146],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003379432,0.000021981998,0.00020312496,0.000057945068,0.000009194061,0.000049631628,0.0002014028,0.0065917764,0.00033228865,0.96022594,0.0014502393,0.030822558],"study_design_scores_gemma":[0.000014140918,0.000031394185,0.00008451061,0.00004440864,0.0000068716045,0.000046485784,0.000029801644,0.00741984,0.00069265626,0.98219246,0.009420521,0.00001700297],"about_ca_topic_score_codex":0.0010749031,"about_ca_topic_score_gemma":0.0004961156,"teacher_disagreement_score":0.98789334,"about_ca_system_score_codex":0.0021333923,"about_ca_system_score_gemma":0.0020240007,"threshold_uncertainty_score":0.06402689},"labels":[],"label_agreement":null},{"id":"W1591612907","doi":"10.1007/978-3-642-02050-6_20","title":"Does Requirements Clustering Lead to Modular Design?","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Cluster analysis; Modular design; Modularity (biology); Context (archaeology); Principal (computer security); Requirements analysis; Software engineering; Software requirements; Software; Data mining; Software design; Systems engineering; Software development; Artificial intelligence; Programming language; Engineering","score_opus":0.036375245390722226,"score_gpt":0.28345310454249434,"score_spread":0.2470778591517721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1591612907","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.119162835,0.00038273697,0.8327262,0.001431693,0.000081005295,0.00015322267,0.00008568725,0.0016377498,0.04433883],"genre_scores_gemma":[0.6357508,0.00041380586,0.34980258,0.000395951,0.000060852875,0.00014336812,0.00026648733,0.0007804664,0.012385799],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974075,0.000973635,0.0001265445,0.0005135273,0.0007219023,0.0002569377],"domain_scores_gemma":[0.9846966,0.007737173,0.0015148292,0.0044755153,0.0012629189,0.00031282424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003719543,0.00074740767,0.0004965968,0.0008589677,0.00078004407,0.0013766282,0.0010987829,0.0011140232,0.009337152],"category_scores_gemma":[0.018358111,0.0010734551,0.0012512346,0.0010268551,0.0013507947,0.004276131,0.0015670313,0.0013010139,0.0022433964],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038679206,0.0002802539,0.0072901268,0.0006736724,0.00016431986,0.00057391176,0.0017735073,0.035790276,0.026553351,0.42563114,0.009253541,0.49162912],"study_design_scores_gemma":[0.00016807961,0.00042342744,0.006753713,0.00038424414,0.00026897367,0.0019264111,0.0012981517,0.21185634,0.040361088,0.69102806,0.045443505,0.00008801364],"about_ca_topic_score_codex":0.00059329043,"about_ca_topic_score_gemma":0.0008972098,"teacher_disagreement_score":0.009337152,"about_ca_system_score_codex":0.000691026,"about_ca_system_score_gemma":0.00088101916,"threshold_uncertainty_score":0.031235874},"labels":[],"label_agreement":null},{"id":"W1592922268","doi":"10.1007/978-3-642-13273-5_8","title":"A New Compound Metric for Software Risk Assessment","year":2010,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Metric (unit); Heuristic; Software; Software engineering; Source code; Risk analysis (engineering); Process (computing); Risk-based testing; Point (geometry); Data mining; Software development; Software construction; Artificial intelligence; Engineering; Programming language; Mathematics","score_opus":0.09333478661709127,"score_gpt":0.3987245809411166,"score_spread":0.30538979432402535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1592922268","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006115889,0.0037712462,0.9550963,0.0007260322,0.0008836759,0.00010716316,0.00048699332,0.0004073874,0.03240534],"genre_scores_gemma":[0.14528131,0.0059355963,0.8006063,0.0004996557,0.0018560637,0.00046705487,0.0012727042,0.00052552077,0.043555852],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968719,0.0005767721,0.00025453774,0.0003461976,0.0018696124,0.00008090146],"domain_scores_gemma":[0.9963889,0.0015529826,0.00038532665,0.0004752392,0.0010215191,0.0001761286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002324427,0.00187958,0.0014554363,0.004349061,0.0007167421,0.003330131,0.0013038955,0.0014728141,0.007077732],"category_scores_gemma":[0.009175855,0.000325868,0.0010142052,0.0043618954,0.0014253509,0.0077755493,0.0020159427,0.0022430813,0.0025046696],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010363015,0.00008470234,0.0015329022,0.0003751766,0.0001228052,0.00018827221,0.00013224875,0.036659196,0.0061622593,0.4838952,0.028182844,0.44256082],"study_design_scores_gemma":[0.000019344947,0.0003521541,0.0019455753,0.0001718469,0.00011342214,0.00096845976,0.000086316075,0.20838565,0.004345832,0.68237376,0.10112956,0.00010804892],"about_ca_topic_score_codex":0.0006355676,"about_ca_topic_score_gemma":0.00092613866,"teacher_disagreement_score":0.007077732,"about_ca_system_score_codex":0.0014756916,"about_ca_system_score_gemma":0.0012585459,"threshold_uncertainty_score":0.02367735},"labels":[],"label_agreement":null},{"id":"W1596155997","doi":"10.1007/11767718_40","title":"Ad Hoc Versus Systematic Planning of Software Releases – A Three-Staged Experiment","year":2006,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Intuition; Post hoc; Scope (computer science); Software; Operations research; Software engineering; Engineering","score_opus":0.030732153343814097,"score_gpt":0.2770940612941766,"score_spread":0.24636190795036247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1596155997","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98667854,0.0001322373,0.006321475,0.0001304153,0.0001375747,0.003086665,0.00030154313,0.0002110222,0.0030003602],"genre_scores_gemma":[0.9610951,0.00017499777,0.021000292,0.0003080059,0.0000980637,0.010420655,0.0005220481,0.00019018992,0.0061906516],"study_design_codex":"randomized_trial","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.991268,0.0035363254,0.0012614556,0.0016746632,0.0015243845,0.00073527935],"domain_scores_gemma":[0.8516881,0.12155902,0.0072725727,0.011257388,0.00347597,0.004747096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013153214,0.0019895306,0.002351137,0.0005804382,0.001314327,0.0027384588,0.0043409574,0.0048475,0.009592177],"category_scores_gemma":[0.081270084,0.0019314032,0.0012357516,0.0006465498,0.0022842775,0.0045991824,0.0021474191,0.0045084828,0.0022485596],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.37205487,0.25508034,0.015347056,0.0036983958,0.0011126747,0.0008259354,0.017247345,0.057305254,0.09917089,0.01368398,0.0036909468,0.1607823],"study_design_scores_gemma":[0.11424395,0.57608134,0.033146422,0.00045452893,0.0015983733,0.0004684682,0.003669152,0.16724417,0.052781433,0.03754712,0.011412434,0.0013525418],"about_ca_topic_score_codex":0.0021683045,"about_ca_topic_score_gemma":0.0015017892,"teacher_disagreement_score":0.013153214,"about_ca_system_score_codex":0.0012699062,"about_ca_system_score_gemma":0.004130132,"threshold_uncertainty_score":0.06956166},"labels":[],"label_agreement":null},{"id":"W1597340042","doi":"10.1007/11603023_7","title":"JQuery: A Generic Code Browser with a Declarative Configuration Language","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Plug-in; Programming language; Personalization; Code (set theory); Interface (matter); Software engineering; Scheme (mathematics); Operating system; Declarative programming; Set (abstract data type); Programming paradigm; World Wide Web; Inductive programming","score_opus":0.017328488391410782,"score_gpt":0.2570951167206205,"score_spread":0.23976662832920975,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1597340042","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0044988007,0.00026325366,0.56001574,0.00033086896,0.0001534641,0.00021468554,0.0029366012,0.42091408,0.010672563],"genre_scores_gemma":[0.1073693,0.0010090084,0.41965795,0.0023638455,0.00022144732,0.0007654515,0.024506591,0.39095542,0.053151045],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99840766,0.00018928727,0.00016664842,0.00030303572,0.00076272734,0.00017067851],"domain_scores_gemma":[0.997255,0.0012120582,0.00014678015,0.0007033108,0.00042889145,0.0002538725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019428272,0.0020652816,0.0014310429,0.001730344,0.00080628024,0.0036929776,0.00427176,0.0031033221,0.029008312],"category_scores_gemma":[0.005598642,0.0029266262,0.0016475596,0.0014385668,0.0015123902,0.005186907,0.004185832,0.0037355265,0.020892875],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019641283,0.0005464446,0.0034966434,0.0014539514,0.00021975866,0.0014051677,0.0018392187,0.008162612,0.07534355,0.065120235,0.55507725,0.285371],"study_design_scores_gemma":[0.0010770098,0.00029238744,0.002526748,0.0004947635,0.00027398608,0.0027426805,0.00030599715,0.13251226,0.11802208,0.058699284,0.6825006,0.00055229606],"about_ca_topic_score_codex":0.0035792342,"about_ca_topic_score_gemma":0.0034587595,"teacher_disagreement_score":0.029008312,"about_ca_system_score_codex":0.0007473069,"about_ca_system_score_gemma":0.0013268915,"threshold_uncertainty_score":0.09704244},"labels":[],"label_agreement":null},{"id":"W1597354914","doi":"10.1109/nafips.2004.1336239","title":"Reverse engineering software architecture using rough clusters","year":2004,"lang":"en","type":"article","venue":"IEEE Annual Meeting of the Fuzzy Information, 2004. Processing NAFIPS '04.","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Reverse engineering; Architecture tradeoff analysis method; Architecture; Computer architecture; Software architecture; Reference architecture; Software; Software engineering; Operating system","score_opus":0.010444362581692734,"score_gpt":0.2355774498002982,"score_spread":0.22513308721860548,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1597354914","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006256884,0.00011192392,0.99254864,0.00010592055,0.000012438374,0.00005824489,0.000032534652,0.0001873115,0.000686129],"genre_scores_gemma":[0.11984502,0.00035353156,0.8779125,0.000048675374,0.000027084227,0.0001901254,0.00022239659,0.00008548889,0.0013150895],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99763024,0.00076725945,0.00016480389,0.0003993859,0.0008864339,0.0001517583],"domain_scores_gemma":[0.9951565,0.002367462,0.00056494947,0.00087200775,0.00091488526,0.00012427775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026893662,0.00096431095,0.0013574325,0.0048989947,0.0016337204,0.0033810039,0.0017766849,0.0013699426,0.0014081267],"category_scores_gemma":[0.011063549,0.0010626293,0.0030327302,0.0036278216,0.0015495613,0.00320834,0.0024680968,0.0017571487,0.0006071907],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007925699,0.0000530327,0.0016740641,0.00018472319,0.00016939225,0.00015872795,0.00071773713,0.7221665,0.0025445055,0.101938024,0.0020920653,0.16822195],"study_design_scores_gemma":[0.0000143764255,0.00003887062,0.00047128665,0.000033768963,0.000039695304,0.000069649555,0.00015435464,0.8981148,0.001564831,0.095996164,0.0034652725,0.000036920574],"about_ca_topic_score_codex":0.010834821,"about_ca_topic_score_gemma":0.009200764,"teacher_disagreement_score":0.010834821,"about_ca_system_score_codex":0.002390382,"about_ca_system_score_gemma":0.0021225854,"threshold_uncertainty_score":0.021543503},"labels":[],"label_agreement":null},{"id":"W1597495308","doi":"10.1007/978-3-540-72667-8_5","title":"Empowering Software Maintainers with Semantic Web Technologies","year":2007,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software maintenance; Software engineering; Traceability; Source code; Static program analysis; Semantic Web; Ontology; Software; Semantic technology; Software development; World Wide Web; Information retrieval; Social Semantic Web; Programming language","score_opus":0.015698393594489428,"score_gpt":0.2682209223338169,"score_spread":0.25252252873932746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1597495308","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064688474,0.0010299721,0.8759766,0.005590979,0.00022820561,0.00010665061,0.00009352534,0.010988359,0.041297216],"genre_scores_gemma":[0.4610306,0.0022511887,0.47698525,0.0011205731,0.0003794773,0.00016933378,0.0007062364,0.0026081703,0.05474915],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.997928,0.00050256425,0.00013284093,0.00021311085,0.0010598191,0.00016354398],"domain_scores_gemma":[0.995773,0.0022769046,0.0002763478,0.0009603063,0.000521878,0.00019158253],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003811529,0.00053768855,0.00030152712,0.000963132,0.0006692264,0.0034011123,0.0013993471,0.0014433566,0.0030858433],"category_scores_gemma":[0.008963117,0.00068983727,0.0004801489,0.0010208553,0.00098548,0.016984697,0.0029285664,0.0024509437,0.0012786849],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013831106,0.00044728754,0.0032217316,0.0004692729,0.00005638819,0.00065573846,0.012051798,0.0038610483,0.044294283,0.16138563,0.027709402,0.7457091],"study_design_scores_gemma":[0.000087616616,0.00023780885,0.0024242292,0.00041015903,0.00024481394,0.0019506034,0.0037774201,0.08240968,0.0826505,0.25069553,0.5749986,0.00011305953],"about_ca_topic_score_codex":0.00054733135,"about_ca_topic_score_gemma":0.0013595977,"teacher_disagreement_score":0.003811529,"about_ca_system_score_codex":0.00037949815,"about_ca_system_score_gemma":0.0011217486,"threshold_uncertainty_score":0.020157516},"labels":[],"label_agreement":null},{"id":"W1599004321","doi":"10.1007/11428817_45","title":"Automatic Transition of Natural Language Software Requirements Specification into Formal Presentation","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":91,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software requirements specification; Natural language; Specification language; Formal specification; Software engineering; Programming language; Ambiguity; Software requirements; Object language; Contradiction; Requirements analysis; Formal methods; Programming language specification; Software; Software development; Natural language processing; Software design; Linguistics","score_opus":0.016440878384865253,"score_gpt":0.28118810877321737,"score_spread":0.26474723038835213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1599004321","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020171093,0.000042896456,0.95742023,0.0001603756,0.0000633171,0.00024924337,0.00019848897,0.017554913,0.004139444],"genre_scores_gemma":[0.40469843,0.00011733031,0.58347714,0.000248167,0.000039076305,0.00041226108,0.0013628895,0.0038042585,0.005840552],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.997375,0.0009535046,0.0002313342,0.00040882186,0.0007902125,0.00024112094],"domain_scores_gemma":[0.99162066,0.005699863,0.0003735237,0.0012917952,0.00090001064,0.00011412067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025200383,0.00052015076,0.00052342785,0.0008806142,0.000421418,0.0018908108,0.0010267313,0.0006784694,0.009583556],"category_scores_gemma":[0.009592513,0.00082998414,0.001224852,0.0005304931,0.0011969974,0.0023366716,0.0024349669,0.001634086,0.0030524936],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011817523,0.000559822,0.0024255116,0.0010525222,0.00009902237,0.0010071724,0.004323872,0.04118022,0.14011453,0.21841994,0.01970196,0.5699337],"study_design_scores_gemma":[0.0003481993,0.00037824948,0.0011376394,0.00027625964,0.00013522536,0.00056782993,0.00059326296,0.48283693,0.27272126,0.18253486,0.058335327,0.00013491635],"about_ca_topic_score_codex":0.0014691519,"about_ca_topic_score_gemma":0.0012438252,"teacher_disagreement_score":0.009583556,"about_ca_system_score_codex":0.00076945865,"about_ca_system_score_gemma":0.0012988248,"threshold_uncertainty_score":0.032060206},"labels":[],"label_agreement":null},{"id":"W1599132279","doi":"10.1007/978-1-4615-0429-0_8","title":"Experimenting with Genetic Algorithms to Devise Optimal Integration Test Orders","year":2003,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Dependency (UML); Computer science; Dimension (graph theory); Measure (data warehouse); Coupling (piping); Class (philosophy); Object (grammar); Algorithm; Genetic algorithm; Mathematical optimization; Mathematics; Data mining; Artificial intelligence; Engineering; Machine learning","score_opus":0.018760803261014636,"score_gpt":0.2539612854256481,"score_spread":0.23520048216463346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1599132279","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1879672,0.00043618266,0.79557824,0.00039989475,0.00008039204,0.00016372347,0.00007714567,0.0018124636,0.013484747],"genre_scores_gemma":[0.4081051,0.00016516981,0.5887178,0.00012829182,0.000019606723,0.0000829232,0.00013100781,0.0002389688,0.0024112002],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99907637,0.0003683068,0.00004473014,0.00012530196,0.00029736725,0.00008779311],"domain_scores_gemma":[0.9945222,0.0044484693,0.00020779116,0.00041437076,0.00034151453,0.00006557949],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020531015,0.0009971766,0.0007576192,0.00081551477,0.00038839644,0.001118151,0.0011929083,0.0014142394,0.0026858759],"category_scores_gemma":[0.010602544,0.0005383993,0.0005410211,0.0009038833,0.0008000493,0.0014950124,0.0005329108,0.0015399504,0.00042222688],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002894382,0.0004238985,0.0020017314,0.00011259073,0.000092698916,0.00006412787,0.00016411065,0.70978147,0.007765028,0.017041963,0.0016956984,0.26056722],"study_design_scores_gemma":[0.00005896422,0.00013502542,0.0002080596,0.000013791911,0.000027419443,0.000024554218,0.00003787999,0.9851718,0.0052477913,0.008461303,0.00060494576,0.000008459629],"about_ca_topic_score_codex":0.0031281805,"about_ca_topic_score_gemma":0.0037127605,"teacher_disagreement_score":0.0031281805,"about_ca_system_score_codex":0.0008297576,"about_ca_system_score_gemma":0.0013134386,"threshold_uncertainty_score":0.01085794},"labels":[],"label_agreement":null},{"id":"W1599355180","doi":"10.1109/icse.2012.6227189","title":"Test confessions: A study of testing practices for plug-in systems","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Test strategy; Computer science; Eclipse; Unit testing; Integration testing; Plug-in; Set (abstract data type); Test (biology); Software engineering; White-box testing; Limiting; Key (lock); Engineering; Operating system; Programming language; Software; Software system; Software construction","score_opus":0.1614351616850639,"score_gpt":0.3923715735112349,"score_spread":0.23093641182617103,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1599355180","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99336666,0.00022191917,0.0038683114,0.0008297466,0.000009856486,0.00009444111,0.00001411408,0.000022269925,0.0015725744],"genre_scores_gemma":[0.99655056,0.00022094556,0.002235731,0.0002703098,0.0000074050654,0.00012378047,0.000030228446,0.000026673177,0.00053439493],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.93967146,0.044488087,0.0031283672,0.0023803404,0.008002359,0.0023294545],"domain_scores_gemma":[0.6110714,0.3174481,0.030485444,0.0103236325,0.022805678,0.007865715],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05251157,0.0004739964,0.00071486103,0.0036022384,0.0036334705,0.0056865243,0.002658283,0.0022954415,0.0011170779],"category_scores_gemma":[0.21296181,0.0010164444,0.0003701665,0.0026139973,0.007793914,0.0072597116,0.0051002586,0.0038273586,0.00023021306],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007660723,0.00040580236,0.06105991,0.00024781266,0.000022597676,0.0007688187,0.90346175,0.000122717,0.0018406586,0.0019276539,0.00041639328,0.029649254],"study_design_scores_gemma":[0.00004092734,0.0010705322,0.063339956,0.000620991,0.000031809777,0.0012218403,0.9166684,0.0019607905,0.0017045661,0.002155265,0.011102159,0.00008276228],"about_ca_topic_score_codex":0.0064948564,"about_ca_topic_score_gemma":0.007688317,"teacher_disagreement_score":0.05251157,"about_ca_system_score_codex":0.005694565,"about_ca_system_score_gemma":0.0061245775,"threshold_uncertainty_score":0.27771103},"labels":[],"label_agreement":null},{"id":"W1600715255","doi":"10.1007/bfb0095026","title":"The introduction and evaluation of object orientation in a company developing real-time embedded systems","year":2000,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Terrafix Geosynthetics (Canada)","funders":"Staffordshire University","keywords":"Computer science; Object-orientation; Suite; Process (computing); Adaptation (eye); Context (archaeology); Object (grammar); Product (mathematics); Orientation (vector space); Quality (philosophy); New product development; Process management; Software engineering; Software; Object-oriented programming; Work (physics); Selection (genetic algorithm); Engineering management; Systems engineering; Artificial intelligence; Engineering; Business; Marketing; Programming language","score_opus":0.021293736115135235,"score_gpt":0.28610480164444474,"score_spread":0.2648110655293095,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1600715255","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9583771,0.0007766737,0.007064536,0.00057993946,0.00006997794,0.00018227029,0.00020506905,0.0001550349,0.03258927],"genre_scores_gemma":[0.9653817,0.0009171793,0.014978896,0.00029977754,0.000037368358,0.000105997904,0.00044998014,0.0000928038,0.017736329],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99784553,0.0011079512,0.00010425664,0.00013397928,0.0006151579,0.00019320968],"domain_scores_gemma":[0.9930443,0.004303635,0.0004985635,0.00027317324,0.001311925,0.0005683332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002816975,0.00036110965,0.00019436209,0.00072862965,0.0007249999,0.0011046387,0.00046711514,0.0008983476,0.0019606703],"category_scores_gemma":[0.007774594,0.00024014802,0.00020813617,0.0013760095,0.0006978169,0.0010485346,0.0007166813,0.0006502322,0.0005184292],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034700516,0.0074967165,0.14182921,0.00093445554,0.000041267638,0.0028369327,0.0409711,0.009814323,0.03486361,0.017591128,0.010325508,0.7298257],"study_design_scores_gemma":[0.0005676704,0.022831619,0.7119114,0.00056172285,0.0002323309,0.003117073,0.042870123,0.026881102,0.07486428,0.008330145,0.1074819,0.0003505706],"about_ca_topic_score_codex":0.010424559,"about_ca_topic_score_gemma":0.014728912,"teacher_disagreement_score":0.010424559,"about_ca_system_score_codex":0.001753297,"about_ca_system_score_gemma":0.0010629436,"threshold_uncertainty_score":0.020727754},"labels":[],"label_agreement":null},{"id":"W1600717382","doi":"10.1007/s10664-012-9200-5","title":"Understanding Ajax applications by connecting client and server-side execution traces","year":2012,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ajax; Computer science; World Wide Web; Web application; Client-side; Web 2.0; Web-based simulation; Field (mathematics); Web page; Web API","score_opus":0.06906239645443948,"score_gpt":0.29564332158658657,"score_spread":0.2265809251321471,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1600717382","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.52774197,0.00026890397,0.4459553,0.00042606535,0.000027847409,0.00048933097,0.0008011542,0.017403172,0.0068861833],"genre_scores_gemma":[0.75537395,0.00037758143,0.2381664,0.00008436743,0.000016783322,0.00017732986,0.0014063946,0.0013024167,0.0030948196],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99872357,0.00033905936,0.00012906826,0.00022343828,0.00051178125,0.000073032905],"domain_scores_gemma":[0.9836246,0.011693351,0.0012971299,0.0011299528,0.0020079059,0.00024707618],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021071695,0.0009374745,0.00031192513,0.002182763,0.00051316834,0.0021191684,0.00093985564,0.0006956467,0.0021296649],"category_scores_gemma":[0.02048454,0.00052726077,0.00024108126,0.000910001,0.0004741037,0.0051003313,0.0012086228,0.0011763738,0.00052880164],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011904283,0.0012526904,0.10450852,0.0012554155,0.00013890896,0.0017683078,0.048650954,0.03183965,0.13438234,0.014580723,0.0063667493,0.6540654],"study_design_scores_gemma":[0.000110487716,0.00086862606,0.12060619,0.00047303573,0.00019006382,0.0018883122,0.010508142,0.67108214,0.1166242,0.032600842,0.044753972,0.00029398527],"about_ca_topic_score_codex":0.0053046946,"about_ca_topic_score_gemma":0.0067311362,"teacher_disagreement_score":0.0053046946,"about_ca_system_score_codex":0.00046218437,"about_ca_system_score_gemma":0.0010349137,"threshold_uncertainty_score":0.011143923},"labels":[],"label_agreement":null},{"id":"W1600856637","doi":"10.1007/978-3-540-79187-4_20","title":"An Industrial Study Using UML Design Metrics for Web Applications","year":2008,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bell (Canada)","funders":"","keywords":"Maintainability; Computer science; Web modeling; Class diagram; Web application; Unified Modeling Language; Web design; Software engineering; Database; Web page; World Wide Web; Programming language; Software","score_opus":0.43911948603517864,"score_gpt":0.44241337614758913,"score_spread":0.0032938901124104913,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1600856637","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93851066,0.0008211863,0.038512465,0.0004390994,0.00004090433,0.00044288466,0.00030951228,0.00025047813,0.02067286],"genre_scores_gemma":[0.938455,0.00069396873,0.054925747,0.00008743049,0.000020482017,0.00017169755,0.00048216677,0.000100582074,0.005062922],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99471635,0.0031905149,0.000312625,0.00033517767,0.0013047812,0.00014065509],"domain_scores_gemma":[0.9632051,0.027137319,0.0010587255,0.0020278145,0.0061971163,0.0003738869],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006406294,0.00048786198,0.00028932412,0.0031328262,0.001131461,0.000973688,0.00081674976,0.0007266447,0.0017484715],"category_scores_gemma":[0.024535557,0.00037549445,0.00037915335,0.003246971,0.0006679129,0.0016383835,0.0008306171,0.00088457984,0.00032185513],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005313515,0.0089174155,0.1019638,0.001165173,0.000102095,0.0010821323,0.020372957,0.014003386,0.02918027,0.029178036,0.012027232,0.7814762],"study_design_scores_gemma":[0.00052642607,0.019926477,0.3574538,0.0018704657,0.0007333051,0.0037884882,0.0354898,0.2789669,0.13036026,0.018701185,0.15188612,0.00029669685],"about_ca_topic_score_codex":0.005970069,"about_ca_topic_score_gemma":0.010403301,"teacher_disagreement_score":0.006406294,"about_ca_system_score_codex":0.0016036051,"about_ca_system_score_gemma":0.0008699322,"threshold_uncertainty_score":0.033880115},"labels":[],"label_agreement":null},{"id":"W1601145182","doi":"10.1109/qsic.2004.24","title":"Machine-learning techniques for software product quality assessment","year":2004,"lang":"en","type":"article","venue":"International Conference on Quality Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Computer science; Software quality; Predictability; Quality (philosophy); Software measurement; Verification and validation; Software; Software metric; Software sizing; Software construction; Software quality control; Software engineering; Software development; Machine learning; Artificial intelligence; Engineering","score_opus":0.112473588299206,"score_gpt":0.41526073520530193,"score_spread":0.3027871469060959,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1601145182","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029838262,0.0033053092,0.9916312,0.00026545068,0.000044261862,0.000050088387,0.00007233124,0.00047821098,0.0011692133],"genre_scores_gemma":[0.27344325,0.006860206,0.71475863,0.00022223729,0.00040330048,0.00050490606,0.0006218312,0.00010450777,0.0030811513],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99791986,0.0008270075,0.00019129338,0.00024887075,0.0007489946,0.00006395194],"domain_scores_gemma":[0.99527705,0.00351161,0.00032013725,0.00030812295,0.0005467303,0.000036348178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025853305,0.0010137932,0.0010748135,0.002622618,0.00038918084,0.0013197535,0.0014577397,0.0013062267,0.0019052069],"category_scores_gemma":[0.010702584,0.00036079582,0.00083304854,0.0037462586,0.00069092325,0.0016772418,0.0008780767,0.0016861109,0.001174284],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004911553,0.000120006385,0.0018146407,0.00045081327,0.00012461624,0.00007276986,0.00009062705,0.34634632,0.0014850877,0.028806856,0.0029297892,0.61770946],"study_design_scores_gemma":[0.000008629417,0.000035030895,0.0005183589,0.000046701734,0.000016946507,0.000045664547,0.000016308179,0.95088017,0.00093807955,0.04445096,0.0030274459,0.000015673675],"about_ca_topic_score_codex":0.0026078222,"about_ca_topic_score_gemma":0.0018546176,"teacher_disagreement_score":0.002622618,"about_ca_system_score_codex":0.00085182284,"about_ca_system_score_gemma":0.000773435,"threshold_uncertainty_score":0.013672709},"labels":[],"label_agreement":null},{"id":"W1601900113","doi":"10.1007/978-3-642-45005-1_9","title":"Detection of SOA Patterns","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science","score_opus":0.015391891169625232,"score_gpt":0.2412983806696382,"score_spread":0.225906489500013,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1601900113","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7542784,0.0013283378,0.21215092,0.000345794,0.00020313878,0.00035388084,0.004015062,0.0074918615,0.019832643],"genre_scores_gemma":[0.90915257,0.0005689474,0.0806669,0.00006553383,0.000046438414,0.00008535778,0.0036840746,0.0001845211,0.005545784],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994778,0.000032483273,0.000028937657,0.00012960873,0.00023982236,0.000091298854],"domain_scores_gemma":[0.9986802,0.00032679475,0.00019838216,0.0002027047,0.0004626833,0.00012938221],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0003339209,0.0004973159,0.00045284393,0.0035350984,0.0004843758,0.0008623214,0.00053536834,0.00046052752,0.0025549848],"category_scores_gemma":[0.0018830918,0.00020548598,0.00054708624,0.0024767064,0.00024747013,0.00068113866,0.0006686963,0.00037669193,0.0013972624],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009072765,0.00031812474,0.091557175,0.00058415247,0.00015479297,0.0022326494,0.0004896483,0.004238065,0.3359108,0.0062909853,0.007840637,0.54947567],"study_design_scores_gemma":[0.00008121342,0.00078780425,0.1867976,0.00018035984,0.00028014206,0.008278639,0.0014733722,0.39022154,0.352715,0.02278111,0.036252994,0.00015021891],"about_ca_topic_score_codex":0.0027547507,"about_ca_topic_score_gemma":0.0048150048,"teacher_disagreement_score":0.0035350984,"about_ca_system_score_codex":0.00034922812,"about_ca_system_score_gemma":0.0007019677,"threshold_uncertainty_score":0.008547246},"labels":[],"label_agreement":null},{"id":"W1604297592","doi":"10.1109/icre.2000.855602","title":"What do you mean I've been practicing without a license? certification &amp; licensing of requirements engineering professionals","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"License; Certification; Software requirements; Requirements engineering; Computer science; Software engineering; Requirements analysis; Work (physics); Software; Software quality; Engineering management; Engineering; Software development; Software design; Mechanical engineering","score_opus":0.05971812495229556,"score_gpt":0.3477914737599661,"score_spread":0.28807334880767055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1604297592","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2542457,0.009806824,0.0068026246,0.5942878,0.007432267,0.0001097935,0.0003151322,0.00023277992,0.12676701],"genre_scores_gemma":[0.93118787,0.0050132684,0.0014096827,0.04136029,0.0008924848,0.00009603034,0.00015892042,0.00010899891,0.019772384],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9823699,0.006112352,0.0008448104,0.0007690271,0.007482432,0.002421429],"domain_scores_gemma":[0.8981375,0.02647525,0.018339753,0.0036149074,0.027046727,0.026385821],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011594504,0.000283592,0.0006189582,0.0014887871,0.004883417,0.006819635,0.0010376727,0.0040764166,0.011686041],"category_scores_gemma":[0.13191912,0.00025976877,0.0005964169,0.0011908823,0.0076371557,0.009084827,0.0035040495,0.006061715,0.004913419],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025190169,0.0010312614,0.24654666,0.00078732305,0.00013401452,0.0021416645,0.07336115,0.0004567793,0.0016618546,0.04533082,0.31283656,0.31545997],"study_design_scores_gemma":[0.00006560822,0.0006499332,0.25427574,0.003401062,0.0000994528,0.007344778,0.263647,0.0013815862,0.0017286906,0.031489216,0.43549857,0.0004183599],"about_ca_topic_score_codex":0.009570051,"about_ca_topic_score_gemma":0.011172399,"teacher_disagreement_score":0.011686041,"about_ca_system_score_codex":0.003673324,"about_ca_system_score_gemma":0.006162733,"threshold_uncertainty_score":0.061318338},"labels":[],"label_agreement":null},{"id":"W1605776721","doi":"10.1002/9781118959312.ch10","title":"Building Models with Categorical Variables","year":2015,"lang":"en","type":"other","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Categorical variable; Variables; Variable (mathematics); Computer science; Simple (philosophy); Regression analysis; Key (lock); Software; Regression; Data mining; Econometrics; Statistics; Machine learning; Mathematics; Programming language","score_opus":0.027568264317211537,"score_gpt":0.2635847566866648,"score_spread":0.23601649236945327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1605776721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007973019,0.00049560115,0.9777828,0.0010311991,0.00029447765,0.00036329654,0.005034742,0.0019384547,0.0050863954],"genre_scores_gemma":[0.16440338,0.0016133352,0.80106527,0.00066206,0.00043918256,0.003623946,0.017996866,0.0009810303,0.009214915],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99034625,0.0057208515,0.00053797395,0.0016738201,0.0013084549,0.00041270757],"domain_scores_gemma":[0.96187377,0.032292735,0.0011219053,0.0022697565,0.002157439,0.00028449405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011161183,0.002411847,0.0020149874,0.0039899144,0.00128539,0.0047147316,0.0035206622,0.0016484729,0.019631146],"category_scores_gemma":[0.057707496,0.0012693978,0.0055031665,0.0042044595,0.0009942311,0.0053047678,0.0035165711,0.0051718093,0.0061332947],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004167748,0.00050388457,0.03412227,0.0012686269,0.0014868736,0.0006035145,0.0014008919,0.33004376,0.00088075997,0.25504673,0.038137645,0.3360883],"study_design_scores_gemma":[0.000066809334,0.00016985246,0.0025873645,0.00042654117,0.00032995947,0.00018157507,0.00045779516,0.6674566,0.0007766189,0.28838062,0.03907625,0.00009003795],"about_ca_topic_score_codex":0.0078467075,"about_ca_topic_score_gemma":0.008261763,"teacher_disagreement_score":0.019631146,"about_ca_system_score_codex":0.0015917184,"about_ca_system_score_gemma":0.002536613,"threshold_uncertainty_score":0.065672755},"labels":[],"label_agreement":null},{"id":"W1606102547","doi":"10.1007/978-3-540-88030-1_37","title":"Visualizing Software Architectural Design Decisions","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Visualization; Computer science; Listing (finance); Software; Information visualization; Software engineering; Set (abstract data type); Software visualization; Architecture; Software architecture; Data science; Software development; Data mining; Software construction; Programming language","score_opus":0.04639195948408347,"score_gpt":0.2914299093787881,"score_spread":0.24503794989470462,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1606102547","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13088162,0.0023125152,0.7825639,0.0015295859,0.00019318223,0.0001339384,0.0035286658,0.029856153,0.04900041],"genre_scores_gemma":[0.43049878,0.0022036994,0.5399114,0.00011149665,0.00007364786,0.0001300445,0.0047066757,0.004245213,0.018119024],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9996606,0.00010462722,0.000023044762,0.00006538805,0.00011554491,0.000030739408],"domain_scores_gemma":[0.99858,0.00074615446,0.00014442878,0.00022021648,0.00021387167,0.000095394615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00060220575,0.0012551217,0.00042282316,0.0024447355,0.0005525608,0.0036725944,0.0007571295,0.0008534636,0.01589191],"category_scores_gemma":[0.0031628474,0.00078087556,0.0006461897,0.0021235612,0.00031988844,0.0029167149,0.0011297147,0.0012081336,0.0022142355],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038805965,0.00018580494,0.00873447,0.0008471663,0.00009075987,0.0007310416,0.0055665285,0.066957474,0.045553874,0.114374734,0.05420122,0.70236886],"study_design_scores_gemma":[0.000088084664,0.00016716088,0.008065704,0.00051678723,0.00015139017,0.00087538233,0.0026285825,0.57720476,0.03229142,0.1400696,0.2377798,0.00016130423],"about_ca_topic_score_codex":0.0036905904,"about_ca_topic_score_gemma":0.007979543,"teacher_disagreement_score":0.01589191,"about_ca_system_score_codex":0.000615366,"about_ca_system_score_gemma":0.00062237727,"threshold_uncertainty_score":0.053163707},"labels":[],"label_agreement":null},{"id":"W1606631627","doi":"10.48550/arxiv.1507.06934","title":"A Neuro-Fuzzy Model for Function Point Calibration","year":2015,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Bank of Canada","funders":"","keywords":"Computer science; Calibration; Fuzzy logic; Software; Neuro-fuzzy; Artificial neural network; Function point; Data mining; Artificial intelligence; Machine learning; Point (geometry); Function (biology); Fuzzy control system; Software development; Mathematics; Statistics","score_opus":0.11402702621626157,"score_gpt":0.2080516039955312,"score_spread":0.09402457777926962,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1606631627","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01303121,0.00034263614,0.97846794,0.00031737503,0.000088710505,0.000054153898,0.0001738518,0.00033582986,0.0071883616],"genre_scores_gemma":[0.8247489,0.0007103581,0.15662867,0.00015166787,0.00008877177,0.00032595312,0.00035087188,0.00010134208,0.016893474],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995009,0.00011545112,0.000024990899,0.00014937036,0.00016297698,0.00004625187],"domain_scores_gemma":[0.99920124,0.00031757995,0.00008371152,0.00007777786,0.0002958153,0.000023925291],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010876926,0.0007937375,0.00080378755,0.00086507795,0.0006866169,0.0017637039,0.0019529585,0.002086518,0.004312484],"category_scores_gemma":[0.003615795,0.0004754478,0.0009168927,0.0011675329,0.00074881787,0.0015747381,0.0006931442,0.0020570613,0.0014283321],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003586618,0.00002312735,0.0004275768,0.000039242062,0.000024626497,0.000057342033,0.000059821818,0.96313816,0.0010380574,0.012546507,0.00077765185,0.021832021],"study_design_scores_gemma":[0.000002594272,0.000009928439,0.00016096808,0.0000084007615,0.000005268435,0.000017359735,0.0000053417025,0.9952193,0.00024552707,0.0036418133,0.0006756114,0.000007807456],"about_ca_topic_score_codex":0.021619791,"about_ca_topic_score_gemma":0.013371104,"teacher_disagreement_score":0.021619791,"about_ca_system_score_codex":0.0015731443,"about_ca_system_score_gemma":0.0011873574,"threshold_uncertainty_score":0.042987883},"labels":[],"label_agreement":null},{"id":"W1607081849","doi":"","title":"A cognitive and user centric based approach for reverse engineering tool design","year":2000,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Reverse engineering; Software engineering; Context (archaeology); Set (abstract data type); Visualization; Software maintenance; Cognition; Systems engineering; Software; Human–computer interaction; Software system; Engineering; Artificial intelligence; Programming language","score_opus":0.09761189240360622,"score_gpt":0.3567598523529033,"score_spread":0.25914795994929707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1607081849","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050196066,0.00007579246,0.9896833,0.0004019515,0.000019577168,0.0003380382,0.000022925338,0.0005884248,0.0038504647],"genre_scores_gemma":[0.060020406,0.000071737915,0.93692833,0.00018320962,0.0000147297105,0.0005421082,0.00006326485,0.00010258973,0.0020736647],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98855865,0.00520847,0.00072171417,0.0014901577,0.0034227031,0.0005983486],"domain_scores_gemma":[0.9869435,0.005564357,0.0007795352,0.0024960595,0.0036047339,0.00061172125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0094038835,0.0022103123,0.00078230037,0.003615443,0.0020455662,0.0071476507,0.0042841034,0.0032589578,0.0025992563],"category_scores_gemma":[0.014735991,0.0015650704,0.0023712644,0.001251647,0.0044956324,0.0045027295,0.0036898945,0.003816415,0.00080547284],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050015625,0.0016420214,0.00676293,0.0019022474,0.0005534694,0.0010744019,0.0188721,0.036298096,0.091396764,0.30389324,0.006536092,0.5305684],"study_design_scores_gemma":[0.00041328146,0.0015748511,0.005750576,0.00083724986,0.00083786173,0.0032769279,0.005018873,0.45492175,0.07609076,0.28972566,0.16092128,0.00063095836],"about_ca_topic_score_codex":0.003612467,"about_ca_topic_score_gemma":0.008730747,"teacher_disagreement_score":0.0094038835,"about_ca_system_score_codex":0.0029651935,"about_ca_system_score_gemma":0.0054872213,"threshold_uncertainty_score":0.049733102},"labels":[],"label_agreement":null},{"id":"W1607292358","doi":"10.1007/978-3-642-14192-8_15","title":"On the Perception of Software Quality Requirements during the Project Lifecycle","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Quality (philosophy); Perception; Software quality; Computer science; Application lifecycle management; Process management; Software engineering; Systems engineering; Software; Business; Engineering; Software development; Psychology; Programming language; Neuroscience; Philosophy","score_opus":0.03876068141847252,"score_gpt":0.3105810167737776,"score_spread":0.2718203353553051,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1607292358","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9912549,0.0001065086,0.0020885493,0.00024050822,0.000005626213,0.000015574018,0.000018096493,0.000009259961,0.00626087],"genre_scores_gemma":[0.99866474,0.000069208465,0.00071094715,0.000028018456,0.0000028800387,0.000009418506,0.000030568706,0.000008458933,0.00047570842],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99753463,0.001149204,0.00010433076,0.00013266856,0.0008620555,0.00021715136],"domain_scores_gemma":[0.95387757,0.034570735,0.0056730234,0.00074881443,0.0034260852,0.0017037517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043999394,0.00016946763,0.00017139544,0.00072038616,0.00043725647,0.0015979283,0.0005258895,0.00089923386,0.0023388292],"category_scores_gemma":[0.03092973,0.00034488228,0.00029797124,0.000484386,0.00068065553,0.001492065,0.00089875335,0.0012309777,0.00024800748],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001551682,0.001087383,0.67277676,0.0003379206,0.000112306494,0.0010338225,0.15690407,0.0025716054,0.025315434,0.0070568174,0.0022579338,0.12899423],"study_design_scores_gemma":[0.000028235203,0.00056070834,0.9303888,0.00011155489,0.00003657163,0.00043981257,0.054816008,0.0071260235,0.0013405617,0.0024657384,0.0026184942,0.00006749274],"about_ca_topic_score_codex":0.0043361643,"about_ca_topic_score_gemma":0.003886527,"teacher_disagreement_score":0.0043999394,"about_ca_system_score_codex":0.0006934108,"about_ca_system_score_gemma":0.0005694149,"threshold_uncertainty_score":0.023269415},"labels":[],"label_agreement":null},{"id":"W1607839285","doi":"10.1109/ase.2004.58","title":"Refactoring use case models on episodes","year":2004,"lang":"en","type":"article","venue":"Spectrum Research Repository (Concordia University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; Computer science; Software engineering; Property (philosophy); Programming language; Software","score_opus":0.06300003456461517,"score_gpt":0.2805878717707753,"score_spread":0.21758783720616015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1607839285","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021315118,0.00023546905,0.9700401,0.00031796133,0.000059134818,0.0005247838,0.00059490756,0.0029356964,0.003976782],"genre_scores_gemma":[0.19400099,0.0008013666,0.7934265,0.00020700044,0.0000692421,0.0009870618,0.0032895938,0.0011911201,0.0060271337],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99324656,0.0021937857,0.0010131538,0.00093872193,0.0021695448,0.00043822662],"domain_scores_gemma":[0.9786226,0.009632353,0.0017897331,0.0068351217,0.0027307835,0.0003894024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059948945,0.0014921487,0.000907903,0.0029780935,0.00074261084,0.0043810117,0.0028360903,0.0018533425,0.0030964245],"category_scores_gemma":[0.023449315,0.0015828566,0.0031725785,0.0018886423,0.001590498,0.0062132725,0.0035349724,0.0028810701,0.0009717909],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040483932,0.00050146953,0.008295762,0.0008740998,0.0003077744,0.0026723729,0.006037616,0.29175937,0.021525897,0.48196808,0.0085256975,0.17712699],"study_design_scores_gemma":[0.0001523493,0.00023528797,0.0011538285,0.00058678084,0.00024374064,0.00059711566,0.00059617934,0.6749986,0.024072714,0.19159147,0.1056214,0.00015053591],"about_ca_topic_score_codex":0.0076914607,"about_ca_topic_score_gemma":0.007342281,"teacher_disagreement_score":0.0076914607,"about_ca_system_score_codex":0.0022904726,"about_ca_system_score_gemma":0.0027122821,"threshold_uncertainty_score":0.031704426},"labels":[],"label_agreement":null},{"id":"W1608143967","doi":"10.1007/11531142_9","title":"Separation of Concerns with Procedures, Annotations, Advice and Pointcuts","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Advice (programming); Separation of concerns; Metadata; Programming language; Software engineering; Source code; Focus (optics); Aspect-oriented programming; Software; World Wide Web","score_opus":0.012840206856417485,"score_gpt":0.2771170648175793,"score_spread":0.26427685796116185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1608143967","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027602832,0.0023308112,0.9770855,0.0003123996,0.00017331104,0.000045881796,0.00007141122,0.0020372625,0.015183274],"genre_scores_gemma":[0.08922457,0.0042229923,0.85577005,0.0003852131,0.00038027234,0.00017802841,0.00071582454,0.002068134,0.04705494],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99829406,0.00030683118,0.00019714447,0.0002681681,0.0008100324,0.00012366635],"domain_scores_gemma":[0.99852604,0.0006350888,0.00012791135,0.00046829364,0.00018210389,0.00006054217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016586557,0.0019807005,0.0011271208,0.0012188427,0.00078549865,0.0036954593,0.0018586895,0.0017164937,0.0062130815],"category_scores_gemma":[0.0040012035,0.0016125522,0.0019007201,0.0016726857,0.001897824,0.0042939493,0.0026954764,0.0048139542,0.0042703343],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012482918,0.00008336131,0.00025591673,0.00069975253,0.00007335746,0.00026020405,0.00092581904,0.004800382,0.013736698,0.58567953,0.016555397,0.37680477],"study_design_scores_gemma":[0.00007458963,0.000088336106,0.00036894414,0.0003416565,0.00014546685,0.0012411734,0.00009393614,0.021294335,0.018623577,0.73238254,0.22527988,0.00006549812],"about_ca_topic_score_codex":0.0007152727,"about_ca_topic_score_gemma":0.0009733662,"teacher_disagreement_score":0.0062130815,"about_ca_system_score_codex":0.0006812786,"about_ca_system_score_gemma":0.0013123245,"threshold_uncertainty_score":0.020784795},"labels":[],"label_agreement":null},{"id":"W1609394770","doi":"10.1109/qrs.2015.30","title":"How Effective Are Code Coverage Criteria?","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Code coverage; Computer science; Test suite; Control flow; Set (abstract data type); Code (set theory); Fault coverage; Metric (unit); Cover (algebra); Data mining; Test case; Reliability engineering; Programming language; Software; Engineering","score_opus":0.029101078576202522,"score_gpt":0.288739980918633,"score_spread":0.2596389023424305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1609394770","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63309735,0.027438063,0.28824338,0.009169825,0.00064795406,0.00044487193,0.0040177247,0.0058485405,0.031092314],"genre_scores_gemma":[0.93495107,0.0016438573,0.059336815,0.00040921045,0.00022829826,0.0001835611,0.0015718051,0.0007507725,0.0009246274],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9458364,0.018797016,0.003989328,0.0049361186,0.02427536,0.002165806],"domain_scores_gemma":[0.7041369,0.23830707,0.015981408,0.011474204,0.02712052,0.002979924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019681454,0.002005912,0.0021335627,0.012068075,0.000712757,0.0046205968,0.0025107826,0.0034636918,0.002037398],"category_scores_gemma":[0.19837402,0.0006632443,0.0013353175,0.006437529,0.00242787,0.009757574,0.0019228826,0.0013968174,0.0009870803],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009769403,0.000616343,0.25567183,0.0018387,0.001278906,0.000488887,0.0011820212,0.038715716,0.014959542,0.016025074,0.011873319,0.65637267],"study_design_scores_gemma":[0.00048216427,0.0032434354,0.328806,0.0024868108,0.0019097212,0.004376981,0.004889259,0.41446674,0.066597536,0.115073636,0.056921586,0.0007460574],"about_ca_topic_score_codex":0.0030673714,"about_ca_topic_score_gemma":0.0031853744,"teacher_disagreement_score":0.019681454,"about_ca_system_score_codex":0.0017250625,"about_ca_system_score_gemma":0.0013944873,"threshold_uncertainty_score":0.1040867},"labels":[],"label_agreement":null},{"id":"W161357454","doi":"10.15388/infedu.2006.01","title":"Random Factors in IOI 2005 Test Case Scoring","year":2006,"lang":"en","type":"article","venue":"Informatics in Education","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Ranking (information retrieval); Test (biology); Statistics; Olympiad; Rank (graph theory); Confidence interval; Variance (accounting); Population; Computer science; Econometrics; Mathematics; Artificial intelligence; Demography; Mathematics education","score_opus":0.011977123254998535,"score_gpt":0.2779785689758177,"score_spread":0.26600144572081913,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W161357454","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.552215,0.0016321185,0.42169723,0.0013857132,0.00016230311,0.0009542086,0.0007146297,0.0011039127,0.02013483],"genre_scores_gemma":[0.9260932,0.00013244017,0.07182736,0.00014486774,0.000059506463,0.00035784958,0.0005073206,0.00013459251,0.00074288127],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.8007006,0.122100286,0.009252614,0.013204183,0.051529128,0.0032132328],"domain_scores_gemma":[0.3160465,0.58959025,0.034219403,0.033037722,0.024928335,0.0021778597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.14365277,0.0013599544,0.0014397764,0.011256708,0.0012218805,0.0058642626,0.0026582072,0.002326902,0.002258782],"category_scores_gemma":[0.60142547,0.00081497955,0.0014794926,0.007968144,0.005042578,0.0047252686,0.0040863287,0.0026047705,0.00063736894],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011113005,0.00039838688,0.4089805,0.00037773588,0.0008144285,0.00037880224,0.0020239358,0.18658279,0.0015483552,0.064822,0.004374037,0.32858774],"study_design_scores_gemma":[0.00022333684,0.0018289366,0.18406342,0.000358853,0.0003877115,0.00094559416,0.00088102854,0.6942203,0.0059798714,0.10488947,0.0058490317,0.00037257842],"about_ca_topic_score_codex":0.006255355,"about_ca_topic_score_gemma":0.0052417847,"teacher_disagreement_score":0.14365277,"about_ca_system_score_codex":0.0046793125,"about_ca_system_score_gemma":0.0022679619,"threshold_uncertainty_score":0.7597175},"labels":[],"label_agreement":null},{"id":"W162010602","doi":"10.1007/978-1-4471-0719-4_15","title":"Analysis of Software Engineering Data Using Computational Intelligence Techniques","year":2001,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Genetic programming; Artificial neural network; Machine learning; Software development; Software; Artificial intelligence; Backpropagation; Data mining; Field (mathematics); Genetic algorithm","score_opus":0.07804424482319854,"score_gpt":0.3183867658257648,"score_spread":0.24034252100256626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W162010602","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010973605,0.003578418,0.9426147,0.0006002154,0.00016375273,0.00016357026,0.003115592,0.0061983583,0.03259172],"genre_scores_gemma":[0.046835814,0.005015857,0.9176436,0.00014029752,0.00013998825,0.00017141874,0.00778358,0.0009841163,0.021285327],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925834,0.00011977108,0.00006410695,0.00010410915,0.00042804392,0.00002569824],"domain_scores_gemma":[0.9969791,0.0021313766,0.00013464292,0.00036292002,0.00036246277,0.000029500108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008282508,0.0008776781,0.0008740616,0.004939075,0.0003785127,0.002569736,0.0010330756,0.00042404135,0.010996963],"category_scores_gemma":[0.005454252,0.00041116338,0.0009235453,0.005219885,0.00040473163,0.0022935267,0.0007265519,0.0011068127,0.006878756],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034629047,0.00006163724,0.0020680355,0.00053073623,0.000060062535,0.00015102659,0.00026520935,0.0057263346,0.004976008,0.028842147,0.024784213,0.9325],"study_design_scores_gemma":[0.000028473294,0.00012482276,0.013249,0.0006965916,0.00026788944,0.0020375862,0.00093103165,0.22076888,0.03679689,0.31117833,0.4137858,0.00013475723],"about_ca_topic_score_codex":0.0014047647,"about_ca_topic_score_gemma":0.0026355453,"teacher_disagreement_score":0.010996963,"about_ca_system_score_codex":0.0005027305,"about_ca_system_score_gemma":0.0006951029,"threshold_uncertainty_score":0.036788523},"labels":[],"label_agreement":null},{"id":"W16205432","doi":"10.48550/arxiv.1512.00764","title":"Extracting Traceability Information from C# Projects","year":2015,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Traceability; Requirements traceability; Computer science; Software engineering; Software maintenance; Source code; Software; Change impact analysis; Code (set theory); Software development; Programming language","score_opus":0.11528227154745826,"score_gpt":0.21191955702082904,"score_spread":0.09663728547337078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W16205432","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31068632,0.001290048,0.62061656,0.001020326,0.0001297564,0.0007077902,0.035641674,0.018125325,0.0117822485],"genre_scores_gemma":[0.43623993,0.0014137296,0.48344716,0.00012037859,0.00011256035,0.00049577,0.072047316,0.0013613136,0.0047619212],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99663925,0.00030550483,0.0003713569,0.00053762086,0.001971063,0.00017518605],"domain_scores_gemma":[0.9758665,0.007892255,0.0039790124,0.004279936,0.0074394997,0.00054285646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012816467,0.0010643805,0.00043679884,0.013394008,0.0006438359,0.001492116,0.0011228631,0.0010338589,0.0011776227],"category_scores_gemma":[0.023507088,0.0005876012,0.00057058426,0.008118851,0.00035462054,0.0026321947,0.0015307021,0.0011765867,0.0010942838],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026916637,0.0004048233,0.06502783,0.0012076828,0.00012299686,0.0019535646,0.0019179217,0.021416647,0.049581237,0.009653583,0.016458284,0.83198625],"study_design_scores_gemma":[0.000114149894,0.0006498259,0.16923957,0.0010603418,0.00029752645,0.0037640529,0.00186744,0.36729687,0.22534095,0.050891977,0.17904465,0.00043264794],"about_ca_topic_score_codex":0.012192145,"about_ca_topic_score_gemma":0.013401163,"teacher_disagreement_score":0.013394008,"about_ca_system_score_codex":0.00077581895,"about_ca_system_score_gemma":0.0028356977,"threshold_uncertainty_score":0.024242342},"labels":[],"label_agreement":null},{"id":"W1631202982","doi":"10.1109/compsac.2015.256","title":"Adaptive Clustering Techniques for Software Components and Architecture","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Cluster analysis; Computer science; Data mining; Fuzzy clustering; FLAME clustering; Hierarchical clustering; Metric (unit); Software metric; Software; Component (thermodynamics); Correlation clustering; CURE data clustering algorithm; Artificial intelligence; Software system; Software construction; Engineering","score_opus":0.04937350249887535,"score_gpt":0.2831740470574781,"score_spread":0.23380054455860277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1631202982","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002407473,0.00036998006,0.9957254,0.00006017279,0.000026766962,0.000035966536,0.000023924904,0.00033975413,0.0010105331],"genre_scores_gemma":[0.14645445,0.00095815636,0.8487012,0.00007302388,0.00007058856,0.00019074368,0.00023067399,0.00017141666,0.0031497646],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987575,0.0002254971,0.00006647335,0.00026257773,0.00062707486,0.000060863465],"domain_scores_gemma":[0.99882334,0.00036566387,0.00012902122,0.00017459334,0.0004831115,0.000024283418],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010177958,0.000842237,0.0006132199,0.0029116217,0.00084693683,0.0008112866,0.0015701393,0.00088555977,0.0020516885],"category_scores_gemma":[0.003988469,0.000416223,0.0011476562,0.003041348,0.0006788604,0.0012661024,0.0008972936,0.0010719029,0.0008637153],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007785999,0.000047168258,0.0009984012,0.00025087368,0.00014620066,0.00011238559,0.00034950848,0.3765232,0.013681745,0.044546224,0.0045853714,0.55868113],"study_design_scores_gemma":[0.000011010086,0.000035536454,0.0008405845,0.000029460756,0.000028554232,0.00015048278,0.000070737675,0.9588624,0.0042980583,0.025383255,0.0102585815,0.0000312286],"about_ca_topic_score_codex":0.0074998112,"about_ca_topic_score_gemma":0.006805646,"teacher_disagreement_score":0.0074998112,"about_ca_system_score_codex":0.0011475867,"about_ca_system_score_gemma":0.001009675,"threshold_uncertainty_score":0.014912307},"labels":[],"label_agreement":null},{"id":"W1637008434","doi":"","title":"Platform library complexity","year":2010,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Suite; Strengths and weaknesses; Computer science; Java; Software; Software engineering; Operating system","score_opus":0.05152573690452441,"score_gpt":0.2826987968242311,"score_spread":0.23117305991970669,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1637008434","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7217777,0.001542828,0.22136039,0.0006182808,0.0000945873,0.0010263543,0.0045961654,0.007118392,0.04186528],"genre_scores_gemma":[0.8960869,0.00066235365,0.09126409,0.00009284077,0.000044844153,0.00043497432,0.0054473025,0.0007315973,0.0052351537],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9919384,0.0008917722,0.00080565555,0.0005470498,0.005218468,0.0005986918],"domain_scores_gemma":[0.955155,0.017222345,0.008685234,0.004862387,0.012120632,0.0019543397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029276488,0.0011598465,0.00063523906,0.010227878,0.0010448109,0.0040101963,0.0012680533,0.0005856103,0.0040474506],"category_scores_gemma":[0.032198194,0.0005144184,0.0011747417,0.005644828,0.0007969395,0.006352801,0.0025421653,0.0010346402,0.0010493896],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008403731,0.0007112415,0.30917382,0.0020585933,0.0009905868,0.00072129635,0.0021668223,0.080744386,0.05500191,0.03843617,0.015145835,0.49400896],"study_design_scores_gemma":[0.00012425803,0.0020211104,0.3421011,0.00062569923,0.0010104979,0.0036233128,0.0023399307,0.3642264,0.12625666,0.05690495,0.09996845,0.00079764664],"about_ca_topic_score_codex":0.005034224,"about_ca_topic_score_gemma":0.005532158,"teacher_disagreement_score":0.010227878,"about_ca_system_score_codex":0.0020207018,"about_ca_system_score_gemma":0.0024165267,"threshold_uncertainty_score":0.015483081},"labels":[],"label_agreement":null},{"id":"W163880961","doi":"","title":"Application of Execution Pattern Mining and Concept Lattice Analysis on Software Structure Evaluation.","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Unix; Software system; Software; Feature (linguistics); Data mining; Software maintenance; Software evolution; Measure (data warehouse); Software sizing; Software construction; Real-time computing; Programming language","score_opus":0.011758758636586776,"score_gpt":0.27459748952204205,"score_spread":0.2628387308854553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W163880961","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.067534514,0.00065983005,0.9255389,0.00039489765,0.000052154082,0.0006582951,0.0009792573,0.0018561969,0.0023259933],"genre_scores_gemma":[0.1940243,0.0001872524,0.80384344,0.000035209832,0.000016897568,0.0003154931,0.0009779246,0.00006135061,0.00053815],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9947936,0.0023255227,0.00037219448,0.0004980579,0.0018619108,0.00014864776],"domain_scores_gemma":[0.9827827,0.012189239,0.0014271276,0.0012850795,0.0019707382,0.00034511532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048974934,0.0007538067,0.00081563054,0.007461102,0.00072864594,0.0016740981,0.0012867643,0.00066859246,0.0012373533],"category_scores_gemma":[0.027659426,0.00035132238,0.0009772871,0.0047766776,0.0008507659,0.0020650781,0.0012429636,0.0010429765,0.00041190992],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047084034,0.0007231093,0.02622135,0.0008454694,0.00042865126,0.00059222366,0.0012202052,0.07288376,0.010259011,0.02903432,0.003751882,0.8535691],"study_design_scores_gemma":[0.0000653085,0.00015919251,0.0038363552,0.00006763759,0.00006587183,0.00047645913,0.00034450868,0.95333564,0.0072651897,0.030716488,0.0036283399,0.000039125705],"about_ca_topic_score_codex":0.005077565,"about_ca_topic_score_gemma":0.005279493,"teacher_disagreement_score":0.007461102,"about_ca_system_score_codex":0.0009808036,"about_ca_system_score_gemma":0.002052744,"threshold_uncertainty_score":0.025900722},"labels":[],"label_agreement":null},{"id":"W1664372859","doi":"","title":"A multi-perspective software visualization environment","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Visualization; Software visualization; Software; Perspective (graphical); Software engineering; Software system; Human–computer interaction; Programming language; Software construction; Data mining; Artificial intelligence","score_opus":0.015843429325902804,"score_gpt":0.27757107812477094,"score_spread":0.2617276487988681,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1664372859","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0058880122,0.0003735887,0.9668255,0.00025839507,0.000059034206,0.00008741322,0.00039122868,0.020732515,0.0053844377],"genre_scores_gemma":[0.04979651,0.0006721753,0.9411324,0.00015465943,0.000058934438,0.00021642356,0.0010492416,0.002722097,0.0041975677],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99928015,0.000192234,0.000045279845,0.00012925922,0.00028151317,0.000071621704],"domain_scores_gemma":[0.99840564,0.00065371307,0.00008249851,0.0002853936,0.00023041693,0.00034238573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011029032,0.0010295645,0.0007405434,0.0013366828,0.0006254266,0.0022609646,0.0021985967,0.0011333294,0.011310784],"category_scores_gemma":[0.002890114,0.0010232852,0.0011497319,0.00083741505,0.00044671568,0.0034712383,0.0050393473,0.0020842142,0.002803012],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010944773,0.0003810217,0.0033778243,0.00094471074,0.00017523095,0.0018502581,0.0039478703,0.03405722,0.1321675,0.065808065,0.06713655,0.6890592],"study_design_scores_gemma":[0.00055312464,0.0004912515,0.002802085,0.00038847871,0.00013920922,0.0026004398,0.0007173956,0.33904278,0.051492788,0.057084683,0.5442573,0.0004305665],"about_ca_topic_score_codex":0.00095918105,"about_ca_topic_score_gemma":0.0018396449,"teacher_disagreement_score":0.011310784,"about_ca_system_score_codex":0.00031641548,"about_ca_system_score_gemma":0.0008057503,"threshold_uncertainty_score":0.03783834},"labels":[],"label_agreement":null},{"id":"W1666209236","doi":"","title":"On tabular expressions","year":2003,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Semantics (computer science); Programming language; Expression (computer science); Software; Theoretical computer science","score_opus":0.1027603024789078,"score_gpt":0.40279683789261944,"score_spread":0.30003653541371167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1666209236","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035430968,0.0037810896,0.9307941,0.0025774909,0.0010815532,0.00019997948,0.000823717,0.0014258501,0.055773098],"genre_scores_gemma":[0.15823276,0.0138542475,0.7507801,0.0067954734,0.0038518906,0.001268409,0.004292227,0.0034439624,0.057480916],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9908862,0.0032484478,0.00095985713,0.0013760981,0.0026927833,0.0008366743],"domain_scores_gemma":[0.9890866,0.0050515933,0.0007466525,0.0022730012,0.0025221573,0.00031998547],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071138567,0.0019125602,0.0014040151,0.003979005,0.0032226283,0.007770732,0.0033012987,0.002513508,0.019371592],"category_scores_gemma":[0.020345327,0.0012780932,0.0029022195,0.0070184455,0.0074532987,0.027058054,0.0069951685,0.005493952,0.0082555525],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000261573,0.000010195016,0.00008842463,0.0001028101,0.000008338676,0.0000743058,0.0003645983,0.0012836613,0.00042134736,0.9706743,0.007505404,0.019440554],"study_design_scores_gemma":[0.0000128863085,0.000014901621,0.000042654894,0.00013846453,0.000014129725,0.0001564468,0.00012466806,0.0049891765,0.0008733652,0.8967695,0.09683773,0.000026215763],"about_ca_topic_score_codex":0.0036846814,"about_ca_topic_score_gemma":0.0021932037,"teacher_disagreement_score":0.019371592,"about_ca_system_score_codex":0.0031291326,"about_ca_system_score_gemma":0.0021917662,"threshold_uncertainty_score":0.064804435},"labels":[],"label_agreement":null},{"id":"W1671697208","doi":"10.7287/peerj.preprints.1138v2","title":"The charming code that error messages are talking about","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Python (programming language); Cyclomatic complexity; Debugging; Computer science; Programming language; Random testing; Code coverage; Source lines of code; Software quality; Software bug; Software; Syntax error; Source code; Code (set theory); Charm (quantum number); Software metric; Abstract syntax tree; Algorithm; Test case; Software development; Particle physics; Machine learning","score_opus":0.07207317227592544,"score_gpt":0.3135958871748,"score_spread":0.24152271489887459,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1671697208","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29709652,0.0061115306,0.41634974,0.027516391,0.009213745,0.0012314975,0.006745044,0.0630041,0.17273147],"genre_scores_gemma":[0.667053,0.002284665,0.18001258,0.012492163,0.0015998723,0.0007260799,0.0037412816,0.016190862,0.11589956],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99168605,0.002464841,0.00059205794,0.0008665065,0.0039257654,0.00046480598],"domain_scores_gemma":[0.950754,0.019820156,0.010392356,0.007806764,0.010174995,0.001051796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026141615,0.0015769986,0.00058862503,0.0024870678,0.0020595687,0.0028588956,0.000929551,0.0022115104,0.013503694],"category_scores_gemma":[0.046489693,0.00055378984,0.00048277265,0.002005445,0.0034545278,0.0053307414,0.0027597174,0.003051101,0.006618674],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015626523,0.0003344461,0.058407843,0.0029268577,0.00027436137,0.0040674014,0.01627849,0.0043713422,0.047951967,0.1461075,0.21124153,0.5064757],"study_design_scores_gemma":[0.00007202816,0.00040193554,0.02924158,0.0022151973,0.00022702881,0.0064423615,0.0035824536,0.010631853,0.06547484,0.05838846,0.8229501,0.00037212577],"about_ca_topic_score_codex":0.0021918914,"about_ca_topic_score_gemma":0.0018684452,"teacher_disagreement_score":0.013503694,"about_ca_system_score_codex":0.0012060588,"about_ca_system_score_gemma":0.0016376926,"threshold_uncertainty_score":0.04517436},"labels":[],"label_agreement":null},{"id":"W1675491259","doi":"10.1109/qrs.2015.45","title":"An Empirical Study of Highly Impactful Bugs in Mozilla Projects","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Software bug; Computer science; Crash; Software regression; Android (operating system); Software; Software quality; Software development; Operating system","score_opus":0.07485372045911827,"score_gpt":0.37727875570063535,"score_spread":0.3024250352415171,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1675491259","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99878865,0.00013038589,0.0004192559,0.00009395302,0.0000032211442,0.00002551964,0.00018569059,0.000019965008,0.00033336476],"genre_scores_gemma":[0.99897313,0.00008110709,0.00038494356,0.000019282159,0.0000044804387,0.000028998003,0.00030809673,0.000008290923,0.0001916535],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9937763,0.0020335084,0.0007802377,0.00094284065,0.0018808634,0.000586208],"domain_scores_gemma":[0.81038785,0.11667104,0.048927885,0.004785074,0.013739211,0.0054889456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007909892,0.00047620962,0.000370253,0.003004863,0.0006620743,0.0012885052,0.0009484397,0.00091484503,0.0013187835],"category_scores_gemma":[0.088846035,0.00042377305,0.0004373238,0.002663541,0.001167326,0.0023685836,0.0012470316,0.0016449875,0.00038290978],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000074518604,0.00021992301,0.9882768,0.00007667927,0.00003778361,0.00018795753,0.0019693675,0.0005782803,0.00024023837,0.00011508924,0.00045901988,0.0077644954],"study_design_scores_gemma":[0.000010475615,0.00037194416,0.98826385,0.000052032112,0.000023139348,0.00035211776,0.0029533606,0.006636151,0.0003004063,0.00013842914,0.0008751503,0.000022947146],"about_ca_topic_score_codex":0.0063672257,"about_ca_topic_score_gemma":0.00742947,"teacher_disagreement_score":0.007909892,"about_ca_system_score_codex":0.0010004458,"about_ca_system_score_gemma":0.00079190836,"threshold_uncertainty_score":0.04183203},"labels":[],"label_agreement":null},{"id":"W168351275","doi":"","title":"An empirical study of the factors affecting co-change frequency of cloned code","year":2013,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Statistics Canada","funders":"","keywords":"Code refactoring; Code (set theory); Computer science; clone (Java method); Software maintenance; Cloning (programming); Software development; Similarity (geometry); Empirical research; Software; Programming language; Biology; Mathematics; Artificial intelligence; Set (abstract data type); Statistics; Genetics; DNA","score_opus":0.19425569687736105,"score_gpt":0.4554084863565816,"score_spread":0.26115278947922055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W168351275","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9985483,0.00010581952,0.0007229997,0.00003811134,0.0000019608897,0.00002764946,0.000108275664,0.000014170976,0.0004327752],"genre_scores_gemma":[0.9988146,0.00004995937,0.00078690494,0.000007557915,0.0000031307138,0.000020634101,0.00015133312,0.0000059892463,0.00015995096],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99313486,0.002626395,0.0007078784,0.0011251834,0.0020150086,0.00039069797],"domain_scores_gemma":[0.7031962,0.22566314,0.044317383,0.007935166,0.015537416,0.0033506765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004798716,0.00030786867,0.00028642014,0.0021068628,0.00040256546,0.0010452952,0.00069012505,0.0006255503,0.0015258684],"category_scores_gemma":[0.0737712,0.00030518477,0.00044947607,0.0029366505,0.00076485734,0.0015514917,0.0004993634,0.0008205312,0.0002872832],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013857725,0.00022130477,0.98696333,0.0000641327,0.00006756202,0.00012427637,0.00064420217,0.00041439917,0.00080861925,0.00008502931,0.000077232165,0.010391406],"study_design_scores_gemma":[0.000007315434,0.00025591478,0.9952859,0.0000125440965,0.000038926322,0.00021373248,0.0006465438,0.0024607847,0.00074317626,0.00006815344,0.0002575771,0.000009534368],"about_ca_topic_score_codex":0.0028155446,"about_ca_topic_score_gemma":0.0033889585,"teacher_disagreement_score":0.004798716,"about_ca_system_score_codex":0.00083285203,"about_ca_system_score_gemma":0.0010505164,"threshold_uncertainty_score":0.025378287},"labels":[],"label_agreement":null},{"id":"W168615561","doi":"","title":"Life and death of software packages: an evolutionary study of Debian","year":2012,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Popularity; Operating system; Source lines of code; Software evolution; Software bug; Software engineering; Database; Software system; Software construction","score_opus":0.13670887466321954,"score_gpt":0.4052443256467245,"score_spread":0.26853545098350495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W168615561","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9820311,0.0010227113,0.008146254,0.0010210912,0.000010702337,0.000031617634,0.00021699366,0.000037426857,0.0074820993],"genre_scores_gemma":[0.9901456,0.0006419089,0.005958252,0.00008979287,0.00001415493,0.000025643245,0.0002402505,0.00004396317,0.0028404053],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99832875,0.00078372523,0.00006462373,0.0003345754,0.00031457795,0.00017375627],"domain_scores_gemma":[0.9830403,0.0094636595,0.0025653476,0.0010600287,0.0028636132,0.0010070488],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0041684983,0.0002464663,0.00037431717,0.0035421771,0.0014709086,0.0019367326,0.0009170404,0.0005561718,0.0020732947],"category_scores_gemma":[0.024913423,0.00026650252,0.0003437567,0.0027226666,0.0020644108,0.0033200826,0.0015615203,0.0010623788,0.00033180934],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040662035,0.00039876738,0.65637696,0.00024105959,0.0001522993,0.0013075934,0.050159834,0.011564497,0.0067040073,0.08272524,0.003914247,0.18604888],"study_design_scores_gemma":[0.000024424477,0.00052145193,0.8110434,0.00017645428,0.00011665085,0.0025695662,0.025340868,0.060889825,0.0036347378,0.04845341,0.04704544,0.00018387318],"about_ca_topic_score_codex":0.005588139,"about_ca_topic_score_gemma":0.0052784863,"teacher_disagreement_score":0.9964578,"about_ca_system_score_codex":0.0021503426,"about_ca_system_score_gemma":0.00059994456,"threshold_uncertainty_score":0.022045374},"labels":[],"label_agreement":null},{"id":"W1689870062","doi":"10.1109/iwpse.2004.19","title":"Studying the evolution of software systems using evolutionary code extractors","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Source code; Software system; Software engineering; Software; Code (set theory); KPI-driven code analysis; Software development; Software construction; Programming language","score_opus":0.03897183920478301,"score_gpt":0.27819398529890443,"score_spread":0.23922214609412143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1689870062","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64760554,0.0005241201,0.34360668,0.00043922683,0.000030805568,0.0002481777,0.0006430789,0.0010567544,0.0058456],"genre_scores_gemma":[0.56216747,0.00075904466,0.43162376,0.00006820173,0.000015713436,0.00015028761,0.0015266924,0.0003402378,0.0033486006],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990011,0.00039321178,0.0000594383,0.0001207602,0.00037679865,0.00004875339],"domain_scores_gemma":[0.99181515,0.005365008,0.0007121342,0.0010668229,0.0009471206,0.00009366926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014042921,0.00032793675,0.00029121398,0.0025691285,0.00043452778,0.0008793134,0.00054543675,0.0005186136,0.0011560511],"category_scores_gemma":[0.014683028,0.000300551,0.0005039217,0.0024771716,0.0003393304,0.0019389616,0.00039376566,0.0008871859,0.0003755619],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015974864,0.00048665673,0.12939215,0.00069724856,0.0003326415,0.0005369005,0.0030908454,0.07720112,0.07640961,0.015453785,0.0017484007,0.6944909],"study_design_scores_gemma":[0.00007687392,0.00050446193,0.13828692,0.00028178614,0.00033151614,0.0012008848,0.0018168989,0.67308193,0.12822399,0.03007534,0.026021147,0.00009823506],"about_ca_topic_score_codex":0.0015526325,"about_ca_topic_score_gemma":0.0041853273,"teacher_disagreement_score":0.0025691285,"about_ca_system_score_codex":0.00039543948,"about_ca_system_score_gemma":0.00055843045,"threshold_uncertainty_score":0.007426679},"labels":[],"label_agreement":null},{"id":"W1715539198","doi":"10.1109/hicss.1998.648334","title":"A compositional approach to software design","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software engineering; Component (thermodynamics); Component-based software engineering; Software development; Software construction; Software design description; Scope (computer science); Software design; Software; Concreteness; Systems engineering; Programming language; Engineering","score_opus":0.05775938251453814,"score_gpt":0.2478785378884566,"score_spread":0.19011915537391846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1715539198","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009527761,0.00057707424,0.9857107,0.0005777171,0.00007107915,0.00014770127,0.000029816123,0.00027506525,0.0116580315],"genre_scores_gemma":[0.044272073,0.0013076677,0.945297,0.00034404028,0.00010984118,0.0005675368,0.00010466934,0.00016699641,0.00783016],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99560994,0.0018345241,0.00037625802,0.000596488,0.0013530862,0.00022963395],"domain_scores_gemma":[0.9968033,0.0013122562,0.00022549824,0.000937403,0.0005658865,0.00015569494],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042900997,0.0012880462,0.0007527787,0.0021441039,0.0025393711,0.0042481255,0.0025982526,0.0017994291,0.005711624],"category_scores_gemma":[0.006368242,0.0011218842,0.0022647143,0.001501742,0.0077066915,0.0042439178,0.0044608032,0.0035180717,0.002245908],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012352485,0.000028616576,0.00011195053,0.00021479306,0.000027249367,0.000098652316,0.0007752035,0.008036498,0.0017670493,0.9666842,0.0012343955,0.02100907],"study_design_scores_gemma":[0.000028593486,0.000050755596,0.00007694104,0.00011802696,0.0000316326,0.0001878544,0.00014758871,0.025457963,0.0020350982,0.89109623,0.080742925,0.000026365808],"about_ca_topic_score_codex":0.0031073657,"about_ca_topic_score_gemma":0.0049823215,"teacher_disagreement_score":0.005711624,"about_ca_system_score_codex":0.0023683875,"about_ca_system_score_gemma":0.0037023122,"threshold_uncertainty_score":0.022688508},"labels":[],"label_agreement":null},{"id":"W1720038619","doi":"10.1007/3-540-45314-8_12","title":"On the Importance of Inter-scenario Relationships in Hierarchical State Machine Design","year":2001,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Dependency (UML); Computer science; GRASP; Finite-state machine; Focus (optics); State (computer science); Transition (genetics); Unified Modeling Language; Theoretical computer science; Artificial intelligence; Programming language","score_opus":0.03907731184482111,"score_gpt":0.26462271314516483,"score_spread":0.22554540130034373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1720038619","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03573424,0.000692103,0.9453123,0.0009892577,0.000071933246,0.00010400149,0.00005823532,0.00023058803,0.016807295],"genre_scores_gemma":[0.7240149,0.0009976445,0.27149042,0.00023710048,0.00011053391,0.00018417707,0.00015096557,0.00015652072,0.0026576985],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99581003,0.002389057,0.00025319756,0.00031575322,0.0009921816,0.00023979695],"domain_scores_gemma":[0.9593251,0.035720855,0.0010481488,0.0025016346,0.0010463222,0.00035808762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005585752,0.0005690421,0.0005858079,0.00066372927,0.001255222,0.0022100222,0.0013029235,0.001587578,0.0047943774],"category_scores_gemma":[0.03812968,0.0010925339,0.0006809968,0.0010416489,0.0029195105,0.00969059,0.002531434,0.0036381276,0.0003991379],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024705663,0.00013177627,0.002333488,0.0003060199,0.000040923067,0.0001905105,0.0007851364,0.31395575,0.0026201757,0.56666714,0.0020167155,0.11070531],"study_design_scores_gemma":[0.000025773139,0.0001007098,0.000428651,0.000067582994,0.000046356887,0.00007990281,0.00013399248,0.49867767,0.0016762762,0.4957801,0.0029609236,0.000022099537],"about_ca_topic_score_codex":0.0020642686,"about_ca_topic_score_gemma":0.0042819497,"teacher_disagreement_score":0.005585752,"about_ca_system_score_codex":0.0009554652,"about_ca_system_score_gemma":0.0009439785,"threshold_uncertainty_score":0.029540598},"labels":[],"label_agreement":null},{"id":"W1739253933","doi":"10.1002/cjce.22261","title":"How do you write and present research well?","year":2015,"lang":"en","type":"article","venue":"The Canadian Journal of Chemical Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science","score_opus":0.04290437010628004,"score_gpt":0.26458149426186356,"score_spread":0.22167712415558352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1739253933","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019937195,0.01737042,0.012434447,0.8104568,0.13975503,0.0008032376,0.00044248372,0.001644679,0.015099221],"genre_scores_gemma":[0.04709906,0.0671694,0.10086636,0.55602413,0.16368644,0.0094678085,0.0019278075,0.0073763076,0.046382733],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.67658013,0.154805,0.03804849,0.014155297,0.10948625,0.0069248066],"domain_scores_gemma":[0.15778303,0.27833226,0.06986368,0.054739386,0.36689532,0.07238633],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.25670642,0.0022795526,0.005279493,0.010929774,0.009514645,0.05903745,0.0055591622,0.019090544,0.028658321],"category_scores_gemma":[0.74550426,0.0023628124,0.0020949754,0.010049039,0.01811851,0.034635916,0.008266788,0.019219404,0.07431204],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011835407,0.00019472037,0.0017827686,0.002902367,0.00018774088,0.00033714715,0.009113945,0.00013377705,0.00094843813,0.0061439713,0.8235245,0.15461223],"study_design_scores_gemma":[0.00015959248,0.00026550557,0.0025226625,0.0114524625,0.00013557555,0.000844274,0.018783962,0.0004194079,0.0009219049,0.02584482,0.9383732,0.0002766126],"about_ca_topic_score_codex":0.002183348,"about_ca_topic_score_gemma":0.0029964934,"teacher_disagreement_score":0.7432936,"about_ca_system_score_codex":0.0083780205,"about_ca_system_score_gemma":0.044745784,"threshold_uncertainty_score":0.91661334},"labels":[],"label_agreement":null},{"id":"W1741323436","doi":"","title":"Integrating SHriMP with the IBM websphere studio workbench","year":2001,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Workbench; Computer science; IBM; Visualization; Shrimp; Software engineering; Data flow diagram; Porting; Programming language; World Wide Web; Database; Software; Artificial intelligence","score_opus":0.014027858379040853,"score_gpt":0.2489606055455842,"score_spread":0.23493274716654333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1741323436","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041003585,0.00036950584,0.8461007,0.0006022995,0.00018749607,0.00040462977,0.0010635523,0.08411767,0.026150562],"genre_scores_gemma":[0.12950467,0.00060681585,0.8318054,0.00030538294,0.00008568381,0.00062511564,0.0043562558,0.01612757,0.016583135],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979382,0.0004572967,0.00016517627,0.00027652193,0.001042098,0.00012074025],"domain_scores_gemma":[0.9960353,0.0017374168,0.00014743695,0.000992089,0.00066738785,0.0004203352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003310306,0.0011401724,0.0006096625,0.0019969256,0.00050920673,0.0027915796,0.0018512601,0.00048562168,0.009800252],"category_scores_gemma":[0.006925026,0.00069822505,0.00080271444,0.0015357622,0.00049488194,0.003059851,0.0035732228,0.0015279317,0.0038049442],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010152673,0.00044651722,0.0054274644,0.0006823772,0.00012512147,0.00199493,0.0048317956,0.0072902334,0.069571525,0.01421968,0.044249836,0.8501452],"study_design_scores_gemma":[0.0006994567,0.0012446182,0.008816171,0.0006235284,0.0002498474,0.0030139172,0.0021014083,0.14231208,0.11989655,0.033004817,0.687614,0.0004235117],"about_ca_topic_score_codex":0.0015061738,"about_ca_topic_score_gemma":0.002010652,"teacher_disagreement_score":0.009800252,"about_ca_system_score_codex":0.00052820565,"about_ca_system_score_gemma":0.0009405589,"threshold_uncertainty_score":0.032785118},"labels":[],"label_agreement":null},{"id":"W1746353556","doi":"10.1007/978-1-84800-044-5_8","title":"Reporting Experiments in Software Engineering","year":2007,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":330,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Guideline; Computer science; Presentation (obstetrics); Unification; Set (abstract data type); Software; Empirical research; Data science; Management science; Software engineering; Engineering; Medicine; Mathematics; Statistics","score_opus":0.06714879045561481,"score_gpt":0.3130913835188216,"score_spread":0.2459425930632068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1746353556","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054549933,0.016645003,0.83558613,0.012759393,0.0029128473,0.00059401087,0.0023135287,0.015603257,0.10813076],"genre_scores_gemma":[0.105813645,0.012936227,0.77198327,0.00689076,0.0040545054,0.0016942524,0.005049093,0.0036697183,0.08790852],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.92660296,0.04721575,0.0050440105,0.0033246367,0.017098956,0.0007136324],"domain_scores_gemma":[0.70195687,0.23316374,0.009568486,0.03876458,0.015017668,0.0015286803],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.039192896,0.0016605242,0.002059753,0.004862311,0.0010198967,0.006030326,0.0039424133,0.00340172,0.024090666],"category_scores_gemma":[0.17332049,0.0010531624,0.0007654287,0.0047956905,0.0039417115,0.010596458,0.0030738015,0.003187699,0.01385964],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017262973,0.00014550876,0.0016998716,0.0013876595,0.00007811672,0.00008506615,0.000945505,0.0014052136,0.0026758015,0.09947188,0.14587343,0.74605936],"study_design_scores_gemma":[0.00012805723,0.00046051937,0.003915942,0.0027215315,0.00019471334,0.0010137501,0.001182843,0.02127011,0.027602132,0.53961,0.40169325,0.00020714698],"about_ca_topic_score_codex":0.0008099793,"about_ca_topic_score_gemma":0.00090259203,"teacher_disagreement_score":0.9608071,"about_ca_system_score_codex":0.0016552331,"about_ca_system_score_gemma":0.002478447,"threshold_uncertainty_score":0.20727432},"labels":[],"label_agreement":null},{"id":"W1751128976","doi":"","title":"Using the cosmic functional size measurement method (iso 19761) as a software requirements improvement mechanism","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Reliability engineering; Functional requirement; Rework; Computer science; Functional specification; Non-functional requirement; Software; Engineering; Software system; Software engineering; Embedded system; Software construction","score_opus":0.11812397682390374,"score_gpt":0.33565855988554383,"score_spread":0.21753458306164009,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1751128976","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14183007,0.0018649397,0.7987165,0.0010844045,0.00032247967,0.0012345604,0.00069906295,0.004405272,0.04984264],"genre_scores_gemma":[0.35591185,0.0007799225,0.6341718,0.00030351768,0.00009731466,0.0011301499,0.0012061085,0.00042034296,0.005979035],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96848005,0.008392112,0.00095864607,0.0013660068,0.020320632,0.0004825899],"domain_scores_gemma":[0.9635756,0.011215422,0.0057814554,0.005157781,0.013856563,0.0004131073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011912041,0.0011932168,0.0006737441,0.007831079,0.0006634106,0.002180762,0.0013001101,0.0011288583,0.0017049115],"category_scores_gemma":[0.03420315,0.00038650827,0.0009409421,0.003526844,0.0010547307,0.0026464413,0.001240998,0.00074962503,0.0007888659],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031700873,0.00033828116,0.039181285,0.00072140904,0.00017276277,0.00028862542,0.0014389445,0.018653888,0.032785133,0.05446609,0.01232587,0.83931065],"study_design_scores_gemma":[0.00030035453,0.005012432,0.31389117,0.0020051647,0.0007100721,0.004011479,0.0026459508,0.19834661,0.18120049,0.038568564,0.25210512,0.0012024962],"about_ca_topic_score_codex":0.005944889,"about_ca_topic_score_gemma":0.0054720277,"teacher_disagreement_score":0.011912041,"about_ca_system_score_codex":0.002412121,"about_ca_system_score_gemma":0.0029471968,"threshold_uncertainty_score":0.0629977},"labels":[],"label_agreement":null},{"id":"W1754212509","doi":"10.1007/978-3-642-41533-3_12","title":"Automatically Searching for Metamodel Well-Formedness Rules in Examples and Counter-Examples","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Metamodeling; Computer science; Rotation formalisms in three dimensions; Set (abstract data type); Programming language; Domain (mathematical analysis); Theoretical computer science; Artificial intelligence; Software engineering; Mathematics","score_opus":0.03098259292486101,"score_gpt":0.2755012836934368,"score_spread":0.2445186907685758,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1754212509","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22410385,0.00055559824,0.7529223,0.00062341115,0.000120142795,0.00045982236,0.0017626137,0.011593988,0.007858258],"genre_scores_gemma":[0.39052862,0.0002469341,0.6007766,0.00021735838,0.00003762088,0.00018581562,0.0043070694,0.0016692654,0.0020307335],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946032,0.0012622843,0.0006983053,0.0010802313,0.0019944604,0.0003613656],"domain_scores_gemma":[0.95816904,0.031624053,0.0024235453,0.0030541152,0.00428375,0.0004454618],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038548342,0.001324367,0.0015506609,0.0043059825,0.0010160112,0.0028496303,0.0024462314,0.0022040547,0.0058602747],"category_scores_gemma":[0.0364624,0.001456425,0.0021237857,0.0020127315,0.0012282939,0.005953462,0.0025646798,0.0017855257,0.0019860277],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014380825,0.000929684,0.047775026,0.0035579,0.00060640386,0.0054998472,0.0038316073,0.04771731,0.07097129,0.089895286,0.017700665,0.7100769],"study_design_scores_gemma":[0.00020081359,0.00031125522,0.005636185,0.00085783674,0.00063242397,0.0040803086,0.0024623314,0.6868978,0.1435191,0.12304557,0.03215407,0.0002023577],"about_ca_topic_score_codex":0.0015837785,"about_ca_topic_score_gemma":0.0041351304,"teacher_disagreement_score":0.0058602747,"about_ca_system_score_codex":0.00073711417,"about_ca_system_score_gemma":0.0017458043,"threshold_uncertainty_score":0.020386577},"labels":[],"label_agreement":null},{"id":"W1761436590","doi":"","title":"Analysis of Software Measures Using Metrology Concepts - ISO 19761 Case Study.","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Metrology; Software; Software measurement; Computer science; Systems engineering; Reliability engineering; Engineering; Software development; Software quality; Mathematics; Statistics; Programming language","score_opus":0.0519040519451182,"score_gpt":0.34866934047033343,"score_spread":0.2967652885252152,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1761436590","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82722336,0.0035628849,0.1139295,0.0021121125,0.00010130452,0.0015849269,0.0011705279,0.00028098,0.05003446],"genre_scores_gemma":[0.8948781,0.0011248367,0.094744205,0.00010731102,0.000019868497,0.00061796326,0.00071136525,0.00005448552,0.007741776],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98945135,0.004641435,0.0005678065,0.00050372595,0.0044999016,0.0003357909],"domain_scores_gemma":[0.98604417,0.008841113,0.0015960109,0.0008697045,0.0024416964,0.0002073503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007980249,0.00053300196,0.00039172196,0.005003158,0.001803988,0.0016430888,0.0013634488,0.0014344957,0.0021348717],"category_scores_gemma":[0.018114667,0.00029654233,0.0006598466,0.004304887,0.0010482126,0.001915921,0.0010712286,0.0008476158,0.00038426166],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007410051,0.0021819589,0.12600367,0.0033808828,0.00026804302,0.021070845,0.042138755,0.032788742,0.019218238,0.15456851,0.025710583,0.5719288],"study_design_scores_gemma":[0.00024915533,0.0031712775,0.25172672,0.001988828,0.0005268581,0.022107204,0.074847326,0.15469837,0.07822386,0.034175586,0.37786493,0.00041989546],"about_ca_topic_score_codex":0.007585094,"about_ca_topic_score_gemma":0.014820874,"teacher_disagreement_score":0.007980249,"about_ca_system_score_codex":0.0042912997,"about_ca_system_score_gemma":0.0015118207,"threshold_uncertainty_score":0.042204082},"labels":[],"label_agreement":null},{"id":"W1762375712","doi":"10.3968/4845","title":"The Software Reliability Increase Method","year":2014,"lang":"en","type":"article","venue":"Studies in sociology of science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Reliability engineering; Code refactoring; Reliability (semiconductor); Computer science; Statistic; Software quality; Software reliability testing; Software; Software metric; Process (computing); Probabilistic logic; Statistics; Software development; Programming language; Mathematics; Artificial intelligence; Engineering","score_opus":0.038089196453151605,"score_gpt":0.3876887638517571,"score_spread":0.3495995673986055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1762375712","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016491096,0.0009313015,0.957652,0.00051637134,0.00017031237,0.00047516852,0.00034595118,0.0035541623,0.019863702],"genre_scores_gemma":[0.2063957,0.0007779867,0.780409,0.000209229,0.00029045632,0.00089894,0.0004406594,0.00042889465,0.010149134],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9948237,0.0012911665,0.00020147048,0.0011081459,0.0024033382,0.00017208314],"domain_scores_gemma":[0.99140656,0.0037240302,0.00091111794,0.0013299596,0.0024183805,0.00020993456],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.0029381602,0.00090442906,0.0006373203,0.004217624,0.0005333783,0.0015519246,0.0013472037,0.0007496268,0.0056529087],"category_scores_gemma":[0.011199576,0.00041981455,0.0009861176,0.0019680855,0.0008797177,0.0023772053,0.0017503526,0.001585491,0.0022262756],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016306317,0.00019989513,0.009618636,0.0007448398,0.00007724797,0.00019999944,0.00057120447,0.011761062,0.012744197,0.06661958,0.007951507,0.8893487],"study_design_scores_gemma":[0.00029210738,0.002039526,0.039140314,0.00081785675,0.0005652252,0.0052567646,0.0009225223,0.49840897,0.05867971,0.15001695,0.24347867,0.0003813317],"about_ca_topic_score_codex":0.00071212556,"about_ca_topic_score_gemma":0.00068695657,"teacher_disagreement_score":0.9994666,"about_ca_system_score_codex":0.0007771701,"about_ca_system_score_gemma":0.0013693855,"threshold_uncertainty_score":0.018910825},"labels":[],"label_agreement":null},{"id":"W176329199","doi":"10.1007/978-3-319-13835-0_16","title":"Combining Static and Dynamic Impact Analysis for Large-Scale Enterprise Systems","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Static analysis; Computer science; Scalability; False positive paradox; Scale (ratio); Dynamic program analysis; Set (abstract data type); Work (physics); Industrial engineering; Distributed computing; Software; Machine learning; Programming language","score_opus":0.010436081523753863,"score_gpt":0.2736924838424286,"score_spread":0.26325640231867475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W176329199","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07775383,0.002527502,0.88793725,0.00057366246,0.00017207513,0.00018091282,0.0006030922,0.0040672733,0.026184436],"genre_scores_gemma":[0.80795395,0.0013412439,0.18353812,0.000109856875,0.00019826127,0.00013372634,0.0008054984,0.00052029337,0.005399049],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985154,0.00034561596,0.000060120252,0.00013224395,0.00078163453,0.0001649344],"domain_scores_gemma":[0.9954978,0.0028505444,0.00021944194,0.00058647693,0.0007238536,0.000121899524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001808312,0.0014494212,0.001021201,0.0040358384,0.0006409694,0.0023378,0.0013551752,0.00069277437,0.0050050146],"category_scores_gemma":[0.0053967806,0.00061188743,0.0013002426,0.0037531292,0.0006104716,0.0045241024,0.0016888992,0.0013558859,0.000818603],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018193101,0.00031031398,0.009096655,0.0005486722,0.0004048882,0.0002330892,0.00021189351,0.5076101,0.010202929,0.05237418,0.005987271,0.41283807],"study_design_scores_gemma":[0.000008686495,0.000057702928,0.0023567684,0.00003165474,0.00009359895,0.00006145961,0.00009766922,0.92683095,0.0024280339,0.06537154,0.0026340317,0.000027804976],"about_ca_topic_score_codex":0.0050331373,"about_ca_topic_score_gemma":0.0074618473,"teacher_disagreement_score":0.0050331373,"about_ca_system_score_codex":0.0010448627,"about_ca_system_score_gemma":0.0011577234,"threshold_uncertainty_score":0.016743422},"labels":[],"label_agreement":null},{"id":"W1763929460","doi":"10.1016/s0167-6423(02)00058-8","title":"A change impact model for changeability assessment in object-oriented software systems","year":2002,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":110,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Maintainability; Computer science; Software metric; Software maintenance; Change impact analysis; Software; Software system; Software quality; Object-oriented programming; Software sizing; Software development; Software engineering; Object-oriented design; Reliability engineering; Component-based software engineering; Systems engineering; Programming language","score_opus":0.06438464849421065,"score_gpt":0.33700594520257726,"score_spread":0.2726212967083666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1763929460","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031057218,0.0001382677,0.95544267,0.00057533797,0.000049330025,0.00025319832,0.00014674957,0.0007739848,0.011563324],"genre_scores_gemma":[0.82854384,0.00021520258,0.1650412,0.00018171145,0.00006508637,0.00043453302,0.00027065503,0.00014880243,0.0050989264],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968262,0.0008925053,0.00019748056,0.00035576415,0.001447896,0.0002801368],"domain_scores_gemma":[0.99306726,0.0038475145,0.00073323783,0.00084380194,0.0012729879,0.0002351954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003078137,0.0009779414,0.0007405296,0.003709173,0.000769847,0.002205344,0.0019154542,0.001977488,0.004656151],"category_scores_gemma":[0.0144955795,0.0004492531,0.0012928679,0.0020807378,0.0015362098,0.00536985,0.0015096162,0.0018068803,0.0008660705],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021436388,0.000544088,0.008080559,0.00023384704,0.00017819954,0.00037328704,0.00061725965,0.6042904,0.005496221,0.22588715,0.0027113075,0.15137333],"study_design_scores_gemma":[0.000026612825,0.00012806566,0.0012853189,0.000021273312,0.000059432885,0.00010132972,0.00006607852,0.9192698,0.0010864651,0.0763881,0.0015422826,0.000025335064],"about_ca_topic_score_codex":0.005568153,"about_ca_topic_score_gemma":0.0039619445,"teacher_disagreement_score":0.005568153,"about_ca_system_score_codex":0.0016817359,"about_ca_system_score_gemma":0.0012131302,"threshold_uncertainty_score":0.016278923},"labels":[],"label_agreement":null},{"id":"W1769239150","doi":"","title":"Software maintenance knowledge-based system (S3mxpert): web-based software maintenance expert training","year":2006,"lang":"en","type":"article","venue":"International Conference on Web-based Education","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Software engineering; Software maintenance; Computer science; Capability Maturity Model; Team software process; Personal software process; Software development; Task (project management); Software construction; Software peer review; Software system; Software analytics; Software; Web application; World Wide Web; Engineering; Systems engineering; Operating system","score_opus":0.03340592950714632,"score_gpt":0.29920836499754255,"score_spread":0.2658024354903962,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1769239150","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28078488,0.00050876,0.5110234,0.0018116183,0.00021284842,0.003981526,0.008331632,0.10811217,0.08523314],"genre_scores_gemma":[0.48832837,0.0004171909,0.4770765,0.0006787142,0.00008106599,0.0027547379,0.008552531,0.0009124252,0.021198621],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99929,0.00020655533,0.00009049933,0.00012151634,0.00022499678,0.00006637131],"domain_scores_gemma":[0.9952734,0.002324749,0.00040426405,0.0007210313,0.0008377114,0.0004388538],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017202899,0.00040651226,0.00043439795,0.0026865655,0.00054591184,0.0012302697,0.00086005946,0.0010319641,0.009489728],"category_scores_gemma":[0.008495909,0.00027269198,0.0003298532,0.0018342956,0.00021968175,0.0025710368,0.0014573544,0.00082123204,0.0045495722],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006125063,0.0021148706,0.016919795,0.00054094335,0.00005519627,0.00027102826,0.0017079882,0.0033633183,0.009901467,0.0034550186,0.030335791,0.93072206],"study_design_scores_gemma":[0.0017141859,0.0058531216,0.2189691,0.0011217917,0.0004957988,0.0026212656,0.0023709373,0.2925024,0.08801096,0.02559829,0.3601159,0.00062621373],"about_ca_topic_score_codex":0.001651739,"about_ca_topic_score_gemma":0.0020492813,"teacher_disagreement_score":0.009489728,"about_ca_system_score_codex":0.00063493976,"about_ca_system_score_gemma":0.0014192172,"threshold_uncertainty_score":0.03174627},"labels":[],"label_agreement":null},{"id":"W1771373694","doi":"10.1002/spe.390","title":"Schrödinger's token","year":2001,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Security token; Computer science; Compiler; Programming language; Parsing; Programmer; Grammar; Context (archaeology); Syntax; Simple (philosophy); Natural language processing; Artificial intelligence; Linguistics; Operating system","score_opus":0.021275170917820926,"score_gpt":0.309870892575293,"score_spread":0.28859572165747205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1771373694","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0453757,0.0007625668,0.6267667,0.004303195,0.0021750994,0.00016439782,0.0006404998,0.0021447693,0.31766713],"genre_scores_gemma":[0.6354004,0.00082371006,0.22049351,0.0014963261,0.00047807148,0.0003215926,0.0004586082,0.001090306,0.1394375],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992766,0.0001804553,0.00006533097,0.00012622499,0.00024970344,0.00010179276],"domain_scores_gemma":[0.9991178,0.00032963319,0.00005180082,0.00029018373,0.00015026392,0.00006029143],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00093904743,0.0003571509,0.00073247193,0.0008205926,0.001184799,0.0020708828,0.0011010057,0.0011409514,0.023604041],"category_scores_gemma":[0.0028707215,0.0003928791,0.0008292324,0.0008699031,0.0024738736,0.003259503,0.0023563055,0.0019905136,0.005189425],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002374298,0.000003967911,0.000049316757,0.000023880322,0.0000037175314,0.000058262995,0.00008344874,0.0009408015,0.0009025043,0.98820484,0.0022830723,0.0074222796],"study_design_scores_gemma":[0.000021756372,0.000018033777,0.00005378503,0.00002233557,0.000008226057,0.00014097505,0.000048225695,0.014289484,0.0038708718,0.95323855,0.028257638,0.000030004961],"about_ca_topic_score_codex":0.001074097,"about_ca_topic_score_gemma":0.0011580592,"teacher_disagreement_score":0.023604041,"about_ca_system_score_codex":0.0011632894,"about_ca_system_score_gemma":0.001307312,"threshold_uncertainty_score":0.07896334},"labels":[],"label_agreement":null},{"id":"W1773448242","doi":"10.1007/978-3-540-85279-7_20","title":"A Case Study on the Impact of Refactoring on Quality and Productivity in an Agile Team","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":140,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Code refactoring; Agile software development; Computer science; Productivity; Quality (philosophy); Software engineering; Technical debt; Software quality; Software development; Empirical research; Software; Process management; Programming language; Engineering","score_opus":0.08793095420036612,"score_gpt":0.3594127995506426,"score_spread":0.2714818453502765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1773448242","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.994371,0.0000649945,0.002708134,0.00017508195,0.000010800428,0.00009560217,0.000029882005,0.000032809356,0.0025117572],"genre_scores_gemma":[0.9933869,0.00008775013,0.0049003856,0.000043317264,0.0000068832496,0.000052440042,0.00004096341,0.000015268899,0.0014659946],"study_design_codex":"design_other","study_design_gemma":"case_report","domain_scores_codex":[0.9962309,0.002136576,0.00016734637,0.00023055598,0.0008063303,0.00042820268],"domain_scores_gemma":[0.96320057,0.028274836,0.0018653786,0.0019475237,0.0025229692,0.002188712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004386781,0.00048895,0.00038732655,0.0011143459,0.0023295938,0.001302078,0.0016262742,0.0022380212,0.0018267818],"category_scores_gemma":[0.013512615,0.00034500362,0.0006023352,0.0017353211,0.00093445013,0.0011968025,0.0010690205,0.0013161803,0.00032093006],"study_design_candidate":"case_report","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006092777,0.044520482,0.20547703,0.0021726175,0.00050242484,0.0852589,0.13217752,0.051715624,0.06388183,0.013113726,0.007708817,0.38737825],"study_design_scores_gemma":[0.003271946,0.061329495,0.4072885,0.001104136,0.0013151732,0.035734933,0.18822734,0.14813638,0.099384315,0.011918652,0.041627903,0.00066130084],"about_ca_topic_score_codex":0.005524526,"about_ca_topic_score_gemma":0.0089561995,"teacher_disagreement_score":0.005524526,"about_ca_system_score_codex":0.0016075019,"about_ca_system_score_gemma":0.0012691341,"threshold_uncertainty_score":0.023199797},"labels":[],"label_agreement":null},{"id":"W177800519","doi":"","title":"Towards dynamic interprocedural analysis in JVMs","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Call graph; Computer science; Reachability; Programming language; Java; Call stack; Class hierarchy; Static analysis; Profiling (computer programming); Theoretical computer science; Graph; Runtime system; Parallel computing; Distributed computing; Object-oriented programming","score_opus":0.007604102090432722,"score_gpt":0.27497544287229436,"score_spread":0.26737134078186164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W177800519","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058710154,0.00025567532,0.9192383,0.00025315178,0.00004685991,0.00007406996,0.00013932772,0.018791439,0.0024909833],"genre_scores_gemma":[0.46981928,0.00028108756,0.523477,0.00025616586,0.000061603525,0.00018174514,0.00047390847,0.0036913012,0.0017578809],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99748456,0.0004548819,0.00010986464,0.00029329106,0.0013827458,0.00027455174],"domain_scores_gemma":[0.9963129,0.0011602194,0.0003195823,0.0015162602,0.00059941236,0.000091588416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011492393,0.00078256615,0.00047925473,0.001918149,0.00078170677,0.0018520197,0.0017598004,0.0008013952,0.0010958669],"category_scores_gemma":[0.0060271374,0.0007047719,0.0009539375,0.000931104,0.001420121,0.0032575796,0.0020328357,0.0021430168,0.0004942505],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045371332,0.0003262615,0.01595374,0.00043273254,0.00015259034,0.00062750885,0.0020456149,0.17414682,0.14500567,0.21319623,0.0085371155,0.43912193],"study_design_scores_gemma":[0.000032803764,0.00009353554,0.0034922145,0.000115673705,0.00005738602,0.00020567076,0.00019224158,0.7675864,0.08715945,0.11356327,0.027410874,0.00009048016],"about_ca_topic_score_codex":0.0034582054,"about_ca_topic_score_gemma":0.0036312772,"teacher_disagreement_score":0.0034582054,"about_ca_system_score_codex":0.0013212865,"about_ca_system_score_gemma":0.0021083993,"threshold_uncertainty_score":0.009586632},"labels":[],"label_agreement":null},{"id":"W1779974843","doi":"10.1016/j.diin.2015.05.015","title":"BinComp: A stratified approach to compiler provenance Attribution","year":2015,"lang":"en","type":"article","venue":"Digital Investigation","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Compiler; Computer science; Compiler correctness; Compiler construction; Optimizing compiler; Interprocedural optimization; Programming language; Loop optimization; Semantics (computer science)","score_opus":0.07929534497281193,"score_gpt":0.2663574744014908,"score_spread":0.1870621294286789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1779974843","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009430369,0.000106275955,0.97056323,0.000099348545,0.00003264008,0.0001464798,0.0005304998,0.017847223,0.001243965],"genre_scores_gemma":[0.11099168,0.00012424539,0.88123703,0.00013328009,0.000033192748,0.00017682221,0.0026167943,0.0026028825,0.0020840194],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9948827,0.0011273103,0.00041376223,0.0010203143,0.002221696,0.00033413243],"domain_scores_gemma":[0.98787993,0.0024344719,0.0010668734,0.0063420236,0.001983392,0.00029334743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004095411,0.0009880832,0.0010278784,0.0031626264,0.0015647776,0.0033549725,0.0025609268,0.001257328,0.002226873],"category_scores_gemma":[0.013795319,0.001258986,0.001657175,0.0030049542,0.0015905207,0.0055818623,0.0052186726,0.0023363081,0.0017877247],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013931453,0.0005494907,0.026607214,0.0005731304,0.00042033353,0.0009945959,0.002083255,0.077826664,0.05780361,0.13256426,0.023877366,0.675307],"study_design_scores_gemma":[0.000059685935,0.00015475803,0.0044242805,0.00010422096,0.00011744103,0.0006290384,0.00022803847,0.737279,0.0907713,0.11827691,0.047789857,0.00016545941],"about_ca_topic_score_codex":0.0068059526,"about_ca_topic_score_gemma":0.010207118,"teacher_disagreement_score":0.0068059526,"about_ca_system_score_codex":0.0013894102,"about_ca_system_score_gemma":0.004451753,"threshold_uncertainty_score":0.021658838},"labels":[],"label_agreement":null},{"id":"W1781502158","doi":"10.1109/simsym.2000.844908","title":"PUMP: a program understanding tool for MODSIM programs","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Software engineering; Interface (matter); Key (lock); Context (archaeology); Object-oriented programming; Programming language; Software; User interface; Code (set theory); Presentation (obstetrics); Operating system","score_opus":0.11961097960683913,"score_gpt":0.3094546942436022,"score_spread":0.18984371463676308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1781502158","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003828824,0.00008913636,0.81526524,0.0001705132,0.000026584095,0.00024339226,0.0011876866,0.17507866,0.004109913],"genre_scores_gemma":[0.08412834,0.0005194452,0.8616062,0.00038911347,0.00006537297,0.0009954799,0.0062188404,0.036035236,0.010041964],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989373,0.00022204044,0.00010480518,0.00016160827,0.0004956133,0.00007863744],"domain_scores_gemma":[0.99692994,0.0019505594,0.00024056947,0.00042067122,0.0003686288,0.000089609646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022005914,0.0009798439,0.0007327086,0.0017549236,0.0004816281,0.0018364806,0.0022270277,0.0008974603,0.02285497],"category_scores_gemma":[0.006942082,0.0009590152,0.0007777676,0.001027063,0.000601877,0.0035441874,0.001997877,0.0017073976,0.0057369145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076954707,0.00028272526,0.0041588275,0.0017136621,0.00013890892,0.0008683526,0.0015419775,0.02438978,0.037896052,0.071658656,0.1703553,0.68622625],"study_design_scores_gemma":[0.0003561575,0.00028878928,0.0027818545,0.0004129845,0.00011902713,0.0015693364,0.0003012787,0.42346838,0.079219304,0.045810606,0.44546857,0.00020373776],"about_ca_topic_score_codex":0.0005955452,"about_ca_topic_score_gemma":0.0007126739,"teacher_disagreement_score":0.02285497,"about_ca_system_score_codex":0.00053188595,"about_ca_system_score_gemma":0.0016230637,"threshold_uncertainty_score":0.07645756},"labels":[],"label_agreement":null},{"id":"W1784491506","doi":"","title":"Informing the Design of a Software-Development Assessment Method with the “Coordinative Artifacts” Framework","year":2014,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Quality (philosophy); Software; Artifact (error); Software development; Software quality; Software engineering; Human–computer interaction; Artificial intelligence","score_opus":0.020914473653509385,"score_gpt":0.2783977314982539,"score_spread":0.2574832578447445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1784491506","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060316375,0.000056102002,0.9889613,0.00045641957,0.000022768176,0.0004451454,0.000032319393,0.00026296248,0.003731343],"genre_scores_gemma":[0.08070386,0.00004603581,0.9176041,0.0000736562,0.000009718238,0.00073192234,0.000092488364,0.00008250016,0.00065567036],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9313812,0.053257465,0.0038142926,0.0031430814,0.007461697,0.0009421613],"domain_scores_gemma":[0.90161425,0.07284883,0.0042626252,0.008779258,0.010912576,0.0015825849],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.045132954,0.0011588815,0.00071652967,0.004025895,0.0020131557,0.0071657225,0.0022106185,0.0025310933,0.004932653],"category_scores_gemma":[0.09296695,0.00093194045,0.0010172209,0.0023210777,0.0039972584,0.009143782,0.00555112,0.0023503576,0.0011443725],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003052892,0.000554296,0.009461309,0.0015806259,0.00013959734,0.0003680407,0.01832817,0.016745308,0.019218815,0.4825939,0.003221669,0.44748297],"study_design_scores_gemma":[0.0003478035,0.0008804293,0.006164891,0.002741296,0.00029742997,0.000872299,0.014734493,0.3064197,0.04183947,0.46267214,0.16279595,0.00023416139],"about_ca_topic_score_codex":0.002201579,"about_ca_topic_score_gemma":0.0036050705,"teacher_disagreement_score":0.045132954,"about_ca_system_score_codex":0.0028761495,"about_ca_system_score_gemma":0.0096681025,"threshold_uncertainty_score":0.2386887},"labels":[],"label_agreement":null},{"id":"W1786186632","doi":"10.1007/978-3-319-24285-9_13","title":"Improving the COSMIC Approximate Sizing Using the Fuzzy Logic EPCU Model","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Sizing; Fuzzy logic; Computer science; Software; Computer engineering; Reliability engineering; Industrial engineering; Algorithm; Artificial intelligence; Engineering; Programming language","score_opus":0.045232950141659484,"score_gpt":0.2710880902295699,"score_spread":0.22585514008791044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1786186632","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.050855998,0.00058250444,0.92631125,0.00015333582,0.000089608526,0.000031912354,0.00013669611,0.000636817,0.021201875],"genre_scores_gemma":[0.8485709,0.00028607735,0.14490177,0.000098058605,0.000041951087,0.000049582013,0.00017586064,0.00015918867,0.0057166163],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999734,0.000048578375,0.000009839467,0.000048792226,0.00013001313,0.00002868886],"domain_scores_gemma":[0.99976593,0.00007639387,0.000018665965,0.00005490168,0.00007454695,0.000009569733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00041186792,0.0005609518,0.00080872833,0.00067538925,0.0004625348,0.0011186582,0.0009991175,0.0006110479,0.004854498],"category_scores_gemma":[0.0009821015,0.00025481856,0.00069977285,0.0009031649,0.000363636,0.0011040837,0.00051860395,0.00064565416,0.00047764162],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006712003,0.00002717313,0.00034342377,0.00004925165,0.000026002304,0.000049489136,0.000026547134,0.90400904,0.00359985,0.024971794,0.0012415693,0.06558887],"study_design_scores_gemma":[0.0000027013225,0.000012688832,0.000071343005,0.0000041936128,0.0000067397787,0.000014694183,0.000004111379,0.99351764,0.0005203688,0.0052518714,0.00059057644,0.0000030036954],"about_ca_topic_score_codex":0.009214027,"about_ca_topic_score_gemma":0.009833417,"teacher_disagreement_score":0.009214027,"about_ca_system_score_codex":0.0009884086,"about_ca_system_score_gemma":0.00090029615,"threshold_uncertainty_score":0.01832074},"labels":[],"label_agreement":null},{"id":"W1787880242","doi":"10.1109/sew.2005.4","title":"Decision Support for Software Release Planning - Methods, Tools, and Practical Experience","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Software release life cycle; Decision support system; Software engineering; Software; Systems engineering; Software development; Software construction; Engineering; Artificial intelligence; Programming language","score_opus":0.08535878872875978,"score_gpt":0.43310100983841787,"score_spread":0.3477422211096581,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1787880242","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.065135054,0.0022039362,0.90484744,0.00357536,0.00013054775,0.00033175596,0.00042954693,0.002845092,0.020501336],"genre_scores_gemma":[0.49902058,0.001252068,0.49549022,0.00012975269,0.00008259075,0.00016659034,0.00043201228,0.0003711542,0.003054987],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98482347,0.009771762,0.0010005987,0.0008736204,0.0030063384,0.0005242362],"domain_scores_gemma":[0.92519563,0.06165235,0.0021607156,0.0048611746,0.0046672663,0.0014629313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019478993,0.0018800881,0.0012816321,0.0027626597,0.0015929935,0.007817631,0.0033735295,0.0021008742,0.011598346],"category_scores_gemma":[0.066124104,0.0011799262,0.00091538747,0.0023190295,0.0012641263,0.0083190985,0.0026256756,0.0029330184,0.002306063],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013028238,0.0010503269,0.005235116,0.0008703468,0.0001211585,0.00026751042,0.002273936,0.09523704,0.0028733322,0.061772306,0.012750767,0.8162453],"study_design_scores_gemma":[0.00040224666,0.0006681449,0.0020577132,0.00054287375,0.00010111259,0.0003570275,0.0017578102,0.8019105,0.0074822297,0.1628145,0.021708325,0.00019746425],"about_ca_topic_score_codex":0.0032204587,"about_ca_topic_score_gemma":0.004655413,"teacher_disagreement_score":0.019478993,"about_ca_system_score_codex":0.0017407492,"about_ca_system_score_gemma":0.0026567678,"threshold_uncertainty_score":0.10301596},"labels":[],"label_agreement":null},{"id":"W1791687657","doi":"10.1007/978-3-642-13881-2_7","title":"Automatic Quality Assessment of Source Code Comments: The JavadocMiner","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Source code; Code (set theory); Quality (philosophy); Programming language","score_opus":0.03315713294727574,"score_gpt":0.32769240608061273,"score_spread":0.294535273133337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1791687657","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13342722,0.0017292254,0.595917,0.0005769529,0.00027766,0.0005226511,0.005710514,0.2589088,0.002930041],"genre_scores_gemma":[0.23329836,0.0004891475,0.7273907,0.00026589402,0.00013913965,0.00033342742,0.013941181,0.017268121,0.0068739476],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99086154,0.0014862991,0.00087914895,0.0017521445,0.0047235233,0.00029741536],"domain_scores_gemma":[0.96580017,0.013755072,0.0046813623,0.0048334715,0.010081576,0.00084833347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006138444,0.0014499541,0.0018053666,0.0060282396,0.0007352189,0.0027608855,0.0024461255,0.0015791525,0.0040470394],"category_scores_gemma":[0.029695245,0.0010476465,0.0013104585,0.002267963,0.00060226716,0.002754938,0.0024275284,0.0013856026,0.0043997825],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078193075,0.00037870105,0.017764436,0.0010629167,0.00024498827,0.00027707723,0.0005236449,0.0032966053,0.068643756,0.0013161425,0.037917037,0.8677927],"study_design_scores_gemma":[0.00058169913,0.0009961375,0.051735412,0.0006281332,0.0005085701,0.001934532,0.0006522688,0.5725929,0.30645087,0.0073577156,0.05613269,0.00042898042],"about_ca_topic_score_codex":0.0019693363,"about_ca_topic_score_gemma":0.0034331116,"teacher_disagreement_score":0.006138444,"about_ca_system_score_codex":0.0007226384,"about_ca_system_score_gemma":0.00180077,"threshold_uncertainty_score":0.03246361},"labels":[],"label_agreement":null},{"id":"W1794343297","doi":"10.1145/2699697","title":"aToucan","year":2015,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":98,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Sequence diagram; Use Case Diagram; Class diagram; Traceability; Programming language; Unified Modeling Language; Software engineering; Completeness (order theory); Activity diagram; Class (philosophy); Requirements traceability; Consistency (knowledge bases); Requirements analysis; Artificial intelligence; Software; Requirement","score_opus":0.14914112278278302,"score_gpt":0.34295324255414494,"score_spread":0.19381211977136192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1794343297","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007277597,0.0029493726,0.06000167,0.0053345053,0.0033235564,0.0006968992,0.013328424,0.060694784,0.84639317],"genre_scores_gemma":[0.051556498,0.0028765588,0.055575524,0.003289152,0.00067224365,0.0010976692,0.03215849,0.017586358,0.83518755],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972174,0.0004171599,0.00014443736,0.00077047804,0.0011493491,0.000301141],"domain_scores_gemma":[0.99575496,0.00082780694,0.00023498244,0.0008967222,0.001357749,0.00092776626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020506058,0.0014906581,0.001106242,0.0026961353,0.002429506,0.008143694,0.003230916,0.003007208,0.472743],"category_scores_gemma":[0.0065023564,0.00077793625,0.0010124525,0.0020490629,0.0010032314,0.0047964845,0.005487242,0.0026513769,0.2882946],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010109887,0.0002515735,0.002122747,0.0008856338,0.00005238814,0.0007845222,0.00063026225,0.001369114,0.0067453873,0.042520974,0.5563378,0.38728866],"study_design_scores_gemma":[0.0000796728,0.00007191285,0.0006198965,0.00012577121,0.000022006938,0.00036203253,0.00014112925,0.0017254468,0.002057321,0.0064101894,0.9883504,0.000034367287],"about_ca_topic_score_codex":0.0044008377,"about_ca_topic_score_gemma":0.0051028556,"teacher_disagreement_score":0.472743,"about_ca_system_score_codex":0.0017164267,"about_ca_system_score_gemma":0.0042969105,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W1794751911","doi":"","title":"Deriving high-level abstractions from legacy software using example-driven clustering","year":2011,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Cluster analysis; Program comprehension; Software maintenance; Graph; Programming language; Theoretical computer science; Software; Software system; Software engineering; Data mining; Machine learning","score_opus":0.3262199322302999,"score_gpt":0.39842409146215707,"score_spread":0.07220415923185719,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1794751911","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010179897,0.000028242752,0.9883917,0.00006766896,0.0000052955133,0.000091632006,0.00006040907,0.0004276518,0.0007475382],"genre_scores_gemma":[0.05800641,0.000088799425,0.94064385,0.000021299955,0.0000057347897,0.00010838292,0.00034328556,0.00022700157,0.00055524637],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99802977,0.00059931417,0.00013591582,0.00030951662,0.00080507994,0.00012042255],"domain_scores_gemma":[0.99475515,0.002718653,0.00044435638,0.0012521114,0.0007071162,0.00012261744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002348558,0.0010368703,0.0006447198,0.0017688636,0.0009791479,0.0016909948,0.002422809,0.0011701811,0.0023234775],"category_scores_gemma":[0.0107584,0.0009419408,0.002044487,0.0013276879,0.0019603027,0.002354188,0.0028673618,0.0017798168,0.0006353711],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009331272,0.00012930429,0.0033581038,0.00075358065,0.00012427651,0.0006452981,0.002109114,0.64467543,0.019400436,0.15367891,0.0026482742,0.1723839],"study_design_scores_gemma":[0.000029946044,0.00006068364,0.0004899072,0.0000808473,0.000048121914,0.00022782825,0.000287855,0.87397873,0.011184371,0.10433081,0.009251444,0.000029511362],"about_ca_topic_score_codex":0.0024456251,"about_ca_topic_score_gemma":0.0048810565,"teacher_disagreement_score":0.0024456251,"about_ca_system_score_codex":0.0010387978,"about_ca_system_score_gemma":0.0017977118,"threshold_uncertainty_score":0.012420535},"labels":[],"label_agreement":null},{"id":"W180288257","doi":"","title":"Automatic bug triage using text categorization.","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":404,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Triage; Categorization; Computer science; Open source; Machine learning; Software bug; Supervised learning; Natural language processing; Artificial intelligence; Data science; Software engineering; Software; Programming language","score_opus":0.026611149846284772,"score_gpt":0.28010506827650933,"score_spread":0.25349391843022456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W180288257","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23578675,0.005756501,0.64193285,0.0026685104,0.0013553987,0.0024040989,0.029315392,0.070816405,0.009964063],"genre_scores_gemma":[0.38888445,0.0010294733,0.54118806,0.0005401889,0.0006883656,0.0012138973,0.05811206,0.00065803144,0.007685405],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99724376,0.00062302826,0.00029655444,0.000768988,0.00092758104,0.00014002375],"domain_scores_gemma":[0.984752,0.0066249534,0.003185672,0.001130913,0.0037878118,0.0005187457],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023960525,0.0016689076,0.001140821,0.014034315,0.0009403676,0.0012464207,0.0015393045,0.001965231,0.0025121814],"category_scores_gemma":[0.012574453,0.00036110348,0.00088823936,0.005256148,0.00035481568,0.0025199363,0.0009766249,0.0011609287,0.00480978],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040757522,0.0008647075,0.033284385,0.0009239887,0.00016779211,0.000507135,0.0005021053,0.004398577,0.03587835,0.00090557756,0.054636646,0.8675231],"study_design_scores_gemma":[0.00029359423,0.00076808664,0.11548345,0.0005831176,0.00037597542,0.0023000347,0.0016997273,0.7353321,0.06313224,0.020362847,0.059397407,0.0002713857],"about_ca_topic_score_codex":0.003930127,"about_ca_topic_score_gemma":0.0050222627,"teacher_disagreement_score":0.014034315,"about_ca_system_score_codex":0.00065967685,"about_ca_system_score_gemma":0.0010545785,"threshold_uncertainty_score":0.012671649},"labels":[],"label_agreement":null},{"id":"W1807333353","doi":"10.1109/wpc.1998.693328","title":"Parsing minimization when extracting information from code in the presence of conditional compilation","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Parsing; Source code; Code (set theory); Programming language; Set (abstract data type); Heuristic; Information extraction; Minification; Natural language processing; Artificial intelligence; Information retrieval; Data mining","score_opus":0.0403953958266925,"score_gpt":0.2609324790138389,"score_spread":0.22053708318714643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1807333353","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07441127,0.00024232613,0.9206213,0.0005384673,0.00002826035,0.00010387885,0.00012268365,0.0020996768,0.0018321234],"genre_scores_gemma":[0.22726943,0.00024958834,0.76886225,0.0002565933,0.000038930753,0.00014902989,0.00060285634,0.001253744,0.0013175701],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957445,0.0014684303,0.00031492772,0.0006131085,0.0015356265,0.00032335427],"domain_scores_gemma":[0.96782166,0.025117598,0.0019437952,0.003199634,0.0017167779,0.00020047033],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044883783,0.0011204024,0.0012304608,0.001540548,0.0009786587,0.0019356313,0.0014644844,0.001358973,0.001156766],"category_scores_gemma":[0.03120713,0.0010162939,0.0011319051,0.00222337,0.0023639435,0.0036841545,0.0017022312,0.0017095617,0.00044423746],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009483361,0.00034670404,0.013086087,0.0017526987,0.00029979047,0.0023174156,0.0028359226,0.22715266,0.12815501,0.11686477,0.012815166,0.4934254],"study_design_scores_gemma":[0.000055317727,0.00022037252,0.0048327725,0.000139181,0.00025713618,0.0016412822,0.000547575,0.687754,0.14173262,0.1508515,0.0118393805,0.00012889582],"about_ca_topic_score_codex":0.0014566254,"about_ca_topic_score_gemma":0.0032685706,"teacher_disagreement_score":0.0044883783,"about_ca_system_score_codex":0.00077491085,"about_ca_system_score_gemma":0.0022784413,"threshold_uncertainty_score":0.023737073},"labels":[],"label_agreement":null},{"id":"W1817101438","doi":"10.1007/978-3-642-37057-1_8","title":"Towards Understanding the Behavior of Classes Using Probabilistic Models of Program Inputs","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Probabilistic logic; Set (abstract data type); Variable (mathematics); Coupling (piping); Theoretical computer science; Probability distribution; Algorithm; Programming language; Artificial intelligence; Mathematics","score_opus":0.08997228213455946,"score_gpt":0.31282967490161007,"score_spread":0.2228573927670506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1817101438","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006978229,0.00024048775,0.9898189,0.0003476226,0.000017589233,0.00004823687,0.00024124487,0.000885955,0.0014217352],"genre_scores_gemma":[0.24017549,0.0014006431,0.75002974,0.0003231689,0.00012702441,0.00034229865,0.0012740797,0.0010092487,0.005318263],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99767095,0.0006321574,0.0001733891,0.000648028,0.0006795087,0.00019604436],"domain_scores_gemma":[0.9894406,0.0069895666,0.0008897468,0.0018231348,0.00060000207,0.00025703027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027077417,0.0018199517,0.0014994843,0.002172549,0.00085893617,0.0053702625,0.00450537,0.0029238225,0.0057858424],"category_scores_gemma":[0.014502012,0.0026042378,0.0045217127,0.0018003116,0.0025352465,0.015144141,0.0029486907,0.0064091845,0.0014222737],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011418922,0.0001974493,0.004641186,0.00048983,0.000159569,0.00018540678,0.0011655844,0.29307902,0.00584077,0.579018,0.0036308437,0.11147812],"study_design_scores_gemma":[0.000008805626,0.000015593727,0.0003219547,0.000041832944,0.000037505677,0.000050979896,0.0000558931,0.658796,0.0011524293,0.33669433,0.0028048477,0.000019769514],"about_ca_topic_score_codex":0.007968132,"about_ca_topic_score_gemma":0.010361598,"teacher_disagreement_score":0.007968132,"about_ca_system_score_codex":0.0025074126,"about_ca_system_score_gemma":0.0022137223,"threshold_uncertainty_score":0.019355595},"labels":[],"label_agreement":null},{"id":"W1819759706","doi":"10.1007/978-3-642-28872-2_22","title":"Cohesive and Isolated Development with Branches","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Alberta; McGill University","funders":"","keywords":"Computer science; Branching (polymer chemistry); Software engineering; Source code; Open source; Software; Database; Operating system; World Wide Web","score_opus":0.014175799800239735,"score_gpt":0.22625982145862242,"score_spread":0.2120840216583827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1819759706","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068478286,0.0033010137,0.5727368,0.0012657205,0.00020286241,0.00019458316,0.00007006511,0.0016224382,0.3521282],"genre_scores_gemma":[0.60624754,0.002439892,0.25700244,0.0004080765,0.00012328797,0.00043104638,0.00025482598,0.0013600736,0.13173276],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99855787,0.00044416054,0.000075419666,0.0002337588,0.000516816,0.00017192197],"domain_scores_gemma":[0.9964497,0.0017304588,0.00023437958,0.00097115216,0.000333371,0.00028100185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012013854,0.00061931874,0.00030641217,0.000825081,0.0010057028,0.0026154795,0.0011126652,0.0008033479,0.0076718656],"category_scores_gemma":[0.00537921,0.00076347886,0.00041876035,0.0010448651,0.0024121017,0.004033041,0.004116672,0.0022144734,0.0022572244],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055069835,0.00009035307,0.0018491102,0.00025747294,0.000019452498,0.00036741368,0.005490606,0.0058396505,0.0037289022,0.6881217,0.006077991,0.28810227],"study_design_scores_gemma":[0.000030137258,0.000116795,0.0017464771,0.0003965125,0.000035187688,0.00057288323,0.0011424614,0.017513093,0.0038618078,0.8254237,0.1491356,0.000025360674],"about_ca_topic_score_codex":0.0006382199,"about_ca_topic_score_gemma":0.001302929,"teacher_disagreement_score":0.0076718656,"about_ca_system_score_codex":0.0006152839,"about_ca_system_score_gemma":0.0014045835,"threshold_uncertainty_score":0.025664926},"labels":[],"label_agreement":null},{"id":"W1822879907","doi":"10.1007/11693017_30","title":"The Pervasiveness of Global Data in Evolving Software Systems","year":2006,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Software system; Software engineering; Data science; Programming language","score_opus":0.023651327539422628,"score_gpt":0.2702795646568284,"score_spread":0.24662823711740578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1822879907","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25610447,0.010630883,0.64802694,0.010078791,0.0004969212,0.000092399925,0.0003403932,0.0012246963,0.07300453],"genre_scores_gemma":[0.92527366,0.0039011645,0.060692437,0.00042643206,0.0003640803,0.00006368707,0.00018701593,0.00036601507,0.008725465],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9977648,0.0005142098,0.00015726333,0.00053246674,0.00084537134,0.00018594533],"domain_scores_gemma":[0.98405,0.009587559,0.0009175043,0.0038227744,0.0010898178,0.0005324677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023531155,0.000340436,0.0005671167,0.0019842414,0.0015099259,0.006035341,0.0014725121,0.0019709507,0.0023009705],"category_scores_gemma":[0.019974379,0.0011080698,0.000654886,0.0030583893,0.00710377,0.02174168,0.005095405,0.003734233,0.00042813682],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017213976,0.000046291483,0.0119353775,0.0004484701,0.00006940048,0.0018340668,0.011560621,0.010383084,0.010016207,0.71815914,0.0047795153,0.23059565],"study_design_scores_gemma":[0.000031985932,0.00012398566,0.0075249146,0.00029278427,0.00015821602,0.003606369,0.003666833,0.04449807,0.009291857,0.8274199,0.10330726,0.000077821904],"about_ca_topic_score_codex":0.0012142787,"about_ca_topic_score_gemma":0.0017936435,"teacher_disagreement_score":0.006035341,"about_ca_system_score_codex":0.00084664434,"about_ca_system_score_gemma":0.00086643454,"threshold_uncertainty_score":0.012444615},"labels":[],"label_agreement":null},{"id":"W1825250184","doi":"10.5220/0002996005210527","title":"AUTOMATED THREAT IDENTIFICATION FOR UML","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Unified Modeling Language; Identification (biology); Applications of UML; UML tool; Software engineering; Programming language; Software","score_opus":0.017553758307750106,"score_gpt":0.305862012297015,"score_spread":0.28830825398926485,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1825250184","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009152383,0.00015091067,0.9618162,0.00028975384,0.00003672974,0.00023965286,0.00025116577,0.02580728,0.002255851],"genre_scores_gemma":[0.13104554,0.00019739356,0.8640075,0.00011897641,0.00003507139,0.00023582754,0.00091299793,0.0011554614,0.002291187],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99499315,0.002137556,0.0003711996,0.0006466339,0.0016308938,0.00022054261],"domain_scores_gemma":[0.9885667,0.007187693,0.0012845172,0.0015227211,0.00128938,0.00014906125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002716335,0.0013680226,0.0007920608,0.0046100845,0.0011223877,0.0025556448,0.0010439472,0.0011868497,0.0047999],"category_scores_gemma":[0.018184923,0.00094202964,0.0018023247,0.0010699455,0.0008983698,0.003382005,0.0023274932,0.001684253,0.0027035638],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026212475,0.00022159024,0.0040866495,0.00057309686,0.0001118144,0.000654647,0.0018545352,0.039364267,0.043525837,0.07177383,0.015887203,0.8216845],"study_design_scores_gemma":[0.0000601917,0.00008180171,0.001224478,0.00021345205,0.00007195956,0.000751744,0.00039459325,0.81986904,0.048418168,0.084539115,0.044264354,0.00011111597],"about_ca_topic_score_codex":0.0028442808,"about_ca_topic_score_gemma":0.0034590007,"teacher_disagreement_score":0.0047999,"about_ca_system_score_codex":0.0014152816,"about_ca_system_score_gemma":0.0023683277,"threshold_uncertainty_score":0.016057253},"labels":[],"label_agreement":null},{"id":"W18306927","doi":"","title":"Towards building effective predictive model in software engineering: a bayesian belief network based approach","year":2010,"lang":"en","type":"article","venue":"Health law in Canada","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Bayesian network; Machine learning; Computer science; Artificial intelligence; Graphical model; Data mining; Software; Set (abstract data type)","score_opus":0.008207087007662367,"score_gpt":0.23965572608836094,"score_spread":0.23144863908069857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W18306927","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004587551,0.0008143374,0.9920397,0.0005345118,0.000023522633,0.000039000686,0.00014187438,0.00019795288,0.0016215548],"genre_scores_gemma":[0.49759644,0.0057058954,0.4900689,0.00056498515,0.0002763113,0.00069988205,0.0013862298,0.00016019298,0.003541199],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980876,0.0007133658,0.000116399926,0.00038528902,0.0005657295,0.00013166414],"domain_scores_gemma":[0.99418527,0.004235542,0.00040295767,0.00017513041,0.00089906476,0.00010202337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035288418,0.0013877136,0.0015184294,0.0039226147,0.00080950826,0.0024440715,0.00285399,0.001869511,0.002244736],"category_scores_gemma":[0.011922406,0.0012008528,0.0012618196,0.0032612262,0.0011267606,0.0039934595,0.0017056202,0.0027594955,0.0006749665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041362626,0.000075176125,0.0026875455,0.00021303019,0.000092527494,0.00010331327,0.00019870768,0.88025534,0.00056579674,0.039839063,0.0016548829,0.07427322],"study_design_scores_gemma":[0.000003338337,0.000009070863,0.00020819478,0.000040019546,0.000016599342,0.000015263024,0.000018459295,0.97653145,0.00013113674,0.022287766,0.00072900724,0.000009662074],"about_ca_topic_score_codex":0.018739922,"about_ca_topic_score_gemma":0.015361655,"teacher_disagreement_score":0.018739922,"about_ca_system_score_codex":0.0023442998,"about_ca_system_score_gemma":0.002043533,"threshold_uncertainty_score":0.037261665},"labels":[],"label_agreement":null},{"id":"W1834928324","doi":"10.1007/978-3-642-15585-7_19","title":"An Empirical Evaluation to Study Benefits of Visual versus Textual Test Coverage Information","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Usability; Unit testing; Visualization; Code coverage; Java; Metric (unit); Abstraction; Test (biology); Code (set theory); Empirical research; Information retrieval; Programming language; Test case; Software engineering; Data mining; Human–computer interaction; Software; Machine learning; Statistics","score_opus":0.034239767217925175,"score_gpt":0.3433009493754234,"score_spread":0.3090611821574982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1834928324","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9704021,0.0023464735,0.015614049,0.00046737233,0.00009444125,0.0009258143,0.0021107043,0.00087671645,0.007162358],"genre_scores_gemma":[0.98173356,0.0003048224,0.013669852,0.00014803896,0.00006864416,0.0005535615,0.0016805861,0.00011623779,0.0017246167],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9834731,0.010408026,0.001245475,0.0011623794,0.0034649114,0.00024609943],"domain_scores_gemma":[0.5428955,0.42263204,0.012209883,0.008112701,0.011971199,0.0021787495],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014597111,0.0008576465,0.00066263834,0.0030235932,0.00045868,0.0013392556,0.0012478717,0.0015502239,0.008242652],"category_scores_gemma":[0.15990004,0.00031128936,0.00076503964,0.0019512316,0.0009832615,0.002499985,0.0010403066,0.0007664483,0.0009915989],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.084126435,0.02598352,0.19328564,0.0063279024,0.0019356557,0.00069358904,0.003286511,0.024419248,0.030790113,0.0046289456,0.011451535,0.61307085],"study_design_scores_gemma":[0.023198914,0.1396423,0.4485523,0.0021120945,0.008612836,0.0037603232,0.004367537,0.2730898,0.06649306,0.008209573,0.021405272,0.00055602915],"about_ca_topic_score_codex":0.0015781643,"about_ca_topic_score_gemma":0.0016666895,"teacher_disagreement_score":0.9854029,"about_ca_system_score_codex":0.000922791,"about_ca_system_score_gemma":0.0008064217,"threshold_uncertainty_score":0.07719785},"labels":[],"label_agreement":null},{"id":"W1838805003","doi":"10.1023/a:1011487332587","title":"Modelling the Likelihood of Software Process Improvement: An Exploratory Study","year":2001,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Capability Maturity Model; Computer science; Sample (material); Process (computing); Empirical research; Software; Process management; Data science; Engineering; Mathematics; Statistics","score_opus":0.03668895162732284,"score_gpt":0.2952584062826417,"score_spread":0.2585694546553189,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1838805003","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9691715,0.00015059828,0.028472122,0.00031756787,0.000004874617,0.00014341768,0.00018843522,0.00006961386,0.0014818433],"genre_scores_gemma":[0.99219257,0.000091354406,0.0069696074,0.000019185649,0.000007617621,0.00008732017,0.00019585565,0.000021717184,0.00041489798],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99195707,0.0063294573,0.00027284684,0.0005548014,0.0005304742,0.0003553474],"domain_scores_gemma":[0.36235157,0.6239454,0.006805261,0.003714147,0.002290761,0.0008928178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02029321,0.00085393054,0.00086406345,0.0014615762,0.0007746857,0.0028621953,0.002689255,0.0029585229,0.0045345686],"category_scores_gemma":[0.22296233,0.00082886114,0.0013274931,0.0017754339,0.0010860023,0.004448671,0.0014092382,0.00341611,0.00054186955],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0057912413,0.003632703,0.49722782,0.00063883397,0.00059523305,0.0011176986,0.007888073,0.3939619,0.001902491,0.027187834,0.0010170746,0.05903912],"study_design_scores_gemma":[0.00022915893,0.0016524511,0.037451126,0.00006171294,0.00022626438,0.00045892558,0.0013379672,0.94161344,0.0012433984,0.014914614,0.00073519425,0.00007584092],"about_ca_topic_score_codex":0.006553068,"about_ca_topic_score_gemma":0.0040925974,"teacher_disagreement_score":0.02029321,"about_ca_system_score_codex":0.0015531855,"about_ca_system_score_gemma":0.0015134607,"threshold_uncertainty_score":0.10732204},"labels":[],"label_agreement":null},{"id":"W1838994652","doi":"","title":"Software estimation: universal models or multiple models?","year":2009,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Estimation; Field (mathematics); Context (archaeology); Exploratory research; Software; Data science; Data modeling; Empirical research; Software development; Exploratory data analysis; Management science; Diversity (politics); Industrial engineering; Econometrics; Operations research; Data mining; Software engineering; Systems engineering; Engineering; Mathematics; Statistics","score_opus":0.02047958386138117,"score_gpt":0.23639682857731006,"score_spread":0.21591724471592888,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1838994652","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023263317,0.017867211,0.92665774,0.017312614,0.00042184032,0.00019041469,0.00055881776,0.0010039142,0.012724081],"genre_scores_gemma":[0.69890994,0.019281045,0.26898983,0.003021779,0.002194377,0.0009532529,0.00096322014,0.00053814426,0.005148349],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97583973,0.015272743,0.0011833389,0.003909222,0.0027247078,0.0010704038],"domain_scores_gemma":[0.91798073,0.06536561,0.0064017083,0.0063454234,0.0029358757,0.0009706613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026340479,0.0025154455,0.004523396,0.0032916886,0.0012409473,0.006268402,0.0055916063,0.0037360962,0.0051458273],"category_scores_gemma":[0.09500253,0.0020799292,0.002718245,0.0045055067,0.0049398625,0.020520829,0.0074918056,0.005730377,0.0012779922],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018680705,0.00024076995,0.023171933,0.0016396224,0.0014248688,0.00040187186,0.0019746872,0.061239734,0.00016920478,0.7065598,0.012285955,0.19070475],"study_design_scores_gemma":[0.000043775362,0.00006780546,0.0024277912,0.0005272781,0.00026728777,0.00023449094,0.0005437071,0.123748325,0.00017934022,0.8613559,0.0105292285,0.000075151],"about_ca_topic_score_codex":0.00583047,"about_ca_topic_score_gemma":0.0051743817,"teacher_disagreement_score":0.026340479,"about_ca_system_score_codex":0.0029043215,"about_ca_system_score_gemma":0.002404813,"threshold_uncertainty_score":0.13930345},"labels":[],"label_agreement":null},{"id":"W1839044010","doi":"10.1109/icsm.2004.1357839","title":"Quality-driven object-oriented re-engineering framework","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Center for Advanced Study, University of Illinois at Urbana-Champaign; University of Toronto","keywords":"Maintainability; Computer science; Business process reengineering; Quality (philosophy); Object-oriented programming; Software quality; Software engineering; Process (computing); Interdependence; Systems engineering; Reliability engineering; Software development; Software; Programming language; Engineering; Manufacturing engineering","score_opus":0.021100748812975294,"score_gpt":0.30017761790852504,"score_spread":0.27907686909554974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1839044010","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001542791,0.00019646878,0.9936786,0.00034009755,0.0000374915,0.00013458989,0.0000343777,0.0008113316,0.0032242609],"genre_scores_gemma":[0.055646285,0.00050481874,0.9377092,0.00015110587,0.000050149076,0.0002770421,0.00022664713,0.00022924572,0.005205393],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99600756,0.0008190875,0.00028108727,0.0003859994,0.0022496355,0.00025665903],"domain_scores_gemma":[0.9966229,0.0005560974,0.00037566497,0.0011596986,0.0010638371,0.0002217856],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060349237,0.0010019639,0.0006779488,0.0018193447,0.00096732867,0.003046547,0.0043104864,0.0015523044,0.0017931234],"category_scores_gemma":[0.0050506257,0.00072054955,0.0015366678,0.0012142279,0.0020560073,0.0028089054,0.0028890234,0.0026465936,0.0010848235],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004832361,0.00020648839,0.0009024402,0.00038893102,0.00010481081,0.0007125507,0.0010256802,0.081332445,0.009495422,0.77496994,0.0055037397,0.1253092],"study_design_scores_gemma":[0.000115601266,0.00015011494,0.00069035485,0.0002687961,0.00017733566,0.0010650379,0.00026901127,0.34921393,0.015092246,0.41611266,0.21672295,0.00012202093],"about_ca_topic_score_codex":0.004329214,"about_ca_topic_score_gemma":0.004415212,"teacher_disagreement_score":0.0060349237,"about_ca_system_score_codex":0.0017966056,"about_ca_system_score_gemma":0.002904214,"threshold_uncertainty_score":0.03191614},"labels":[],"label_agreement":null},{"id":"W1844242138","doi":"","title":"SAGA-ML: an active learning system for semi-automated gameplay analysis","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Correctness; Set (abstract data type); Context (archaeology); Black box; Software; Artificial intelligence; Machine learning; Visualization; Human–computer interaction; Component (thermodynamics); Video game; Software engineering; Programming language; Multimedia","score_opus":0.016721533898730493,"score_gpt":0.2915047009522573,"score_spread":0.2747831670535268,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1844242138","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036135016,0.000044977634,0.9498114,0.00006604484,0.000027778631,0.00009911508,0.00024369836,0.045500018,0.0005934979],"genre_scores_gemma":[0.20026965,0.00008110181,0.7921087,0.00020543089,0.000044666594,0.0007568422,0.0012559115,0.0026511701,0.0026265502],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99845576,0.000659908,0.00010663552,0.0002952292,0.00038988804,0.000092582246],"domain_scores_gemma":[0.99433124,0.0039121513,0.0002895144,0.00066783006,0.00057333143,0.00022592452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023467557,0.0014240185,0.0010753517,0.0015016455,0.00056670065,0.001907998,0.0039151064,0.0021882243,0.009640556],"category_scores_gemma":[0.010879669,0.0007499594,0.0007875102,0.00055864337,0.0010887696,0.003088896,0.0025675248,0.0026665295,0.0039319173],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012610198,0.0008670392,0.0052730385,0.0006312201,0.00027683313,0.00033012184,0.0009735419,0.14651372,0.027556209,0.020599062,0.027263002,0.7684551],"study_design_scores_gemma":[0.0000481518,0.000085700565,0.00031815973,0.000019002471,0.00001729327,0.000055203047,0.000024163966,0.97393996,0.00852472,0.011879621,0.005062691,0.000025350417],"about_ca_topic_score_codex":0.0014609318,"about_ca_topic_score_gemma":0.0020575125,"teacher_disagreement_score":0.009640556,"about_ca_system_score_codex":0.0005975817,"about_ca_system_score_gemma":0.000848417,"threshold_uncertainty_score":0.03225082},"labels":[],"label_agreement":null},{"id":"W1845560990","doi":"10.1007/978-3-642-13881-2_8","title":"Towards Approximating COSMIC Functional Size from User Requirements in Agile Development Processes Using Text Mining","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Agile software development; Computer science; User story; Granularity; Scrum; Software; User requirements document; Software engineering; Function point; Functional requirement; Software development; COSMIC cancer database; Requirements analysis; Data mining; Programming language","score_opus":0.04357323396090906,"score_gpt":0.2753789324966259,"score_spread":0.2318056985357168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1845560990","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21789709,0.000482287,0.7761799,0.00016585985,0.000028834656,0.000103305676,0.0011472955,0.0018990191,0.0020964372],"genre_scores_gemma":[0.6203736,0.00026147312,0.37484097,0.00005170183,0.00002906207,0.00022480238,0.0030739638,0.00025005164,0.0008942929],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99829537,0.00062173913,0.00018238263,0.00030668618,0.0005043186,0.000089455534],"domain_scores_gemma":[0.97993374,0.015964681,0.0012952238,0.0011333,0.0015004297,0.00017266176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020425238,0.00086391566,0.0007202956,0.0033731195,0.00040765604,0.001418162,0.0011731085,0.00093770714,0.0009428895],"category_scores_gemma":[0.020347975,0.00051186694,0.0010371647,0.0026952839,0.00041136096,0.0022857157,0.0009659712,0.0012184253,0.00061447953],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038933242,0.0003258983,0.050905947,0.00045045622,0.00015436271,0.00032363326,0.00073053775,0.32906055,0.012010854,0.009492406,0.0027857816,0.59337026],"study_design_scores_gemma":[0.0000059224617,0.000032842734,0.0044275983,0.000029436893,0.000018044762,0.000053342483,0.000088827815,0.98601496,0.0023200384,0.0063379332,0.00066138303,0.000009671735],"about_ca_topic_score_codex":0.005689194,"about_ca_topic_score_gemma":0.005667618,"teacher_disagreement_score":0.005689194,"about_ca_system_score_codex":0.0008539553,"about_ca_system_score_gemma":0.0008884595,"threshold_uncertainty_score":0.011312187},"labels":[],"label_agreement":null},{"id":"W1848096036","doi":"10.1109/metric.2001.915518","title":"A fuzzy logic based set of measures for software project similarity: validation and possible improvements","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Fuzzy logic; Set (abstract data type); Similarity (geometry); Fuzzy set; Software; Data mining; Artificial intelligence; Programming language; Software engineering","score_opus":0.07409619173250517,"score_gpt":0.30586543407461186,"score_spread":0.2317692423421067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1848096036","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08235398,0.00049594487,0.909614,0.00070585747,0.00010528144,0.00044775577,0.0005286878,0.0005691385,0.00517934],"genre_scores_gemma":[0.33629414,0.00020432241,0.6617343,0.000109874825,0.000039734976,0.0004840102,0.0005782985,0.00003830945,0.000517058],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9711592,0.010184121,0.0038680602,0.0028988158,0.011402233,0.00048767024],"domain_scores_gemma":[0.9225327,0.045692746,0.005469918,0.007839085,0.017795946,0.0006696424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025673106,0.000756146,0.0010445511,0.007870646,0.001475839,0.004460401,0.0032789942,0.0017860366,0.001807499],"category_scores_gemma":[0.08552711,0.00041910517,0.0017541009,0.0046606213,0.0021958714,0.0066891722,0.0017964431,0.0020139012,0.00042914812],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010967787,0.0009886468,0.040689267,0.0013228667,0.00053114677,0.00026628462,0.002219101,0.08392181,0.010025756,0.16090882,0.0036907431,0.6943388],"study_design_scores_gemma":[0.00022468127,0.0012636629,0.02533858,0.000398402,0.00032728803,0.00062019273,0.0011300114,0.8585364,0.01721703,0.08701837,0.0076229093,0.0003024204],"about_ca_topic_score_codex":0.0038237898,"about_ca_topic_score_gemma":0.0029429065,"teacher_disagreement_score":0.025673106,"about_ca_system_score_codex":0.0039538257,"about_ca_system_score_gemma":0.0025495964,"threshold_uncertainty_score":0.13577396},"labels":[],"label_agreement":null},{"id":"W1850580823","doi":"","title":"Development of a scaling factors framework to improve the approximation of software functional size with cosmic - iso19761","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software development; Computer science; Software sizing; Software development process; Software metric; Software engineering; Software construction; Personal software process; Software peer review; Software","score_opus":0.016452890903191963,"score_gpt":0.23178274637398288,"score_spread":0.21532985547079092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1850580823","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001720264,0.000426849,0.9951014,0.00008498217,0.000050406576,0.000055906796,0.000065029715,0.0003064259,0.0021888437],"genre_scores_gemma":[0.103716075,0.0014007256,0.8898398,0.00019590341,0.00023549785,0.00052621926,0.0007685114,0.0006028262,0.002714492],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9948131,0.0016231003,0.00038578548,0.00075680035,0.0020969803,0.00032427255],"domain_scores_gemma":[0.99355465,0.0026224835,0.00064748974,0.0005791846,0.002426015,0.00017005991],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064092604,0.0023928657,0.0015536168,0.0066813855,0.0010021509,0.0032685674,0.0027073736,0.0017278115,0.0052781436],"category_scores_gemma":[0.019758005,0.0009270717,0.0028498136,0.0047432994,0.0013523606,0.0040601296,0.0028823977,0.0032623024,0.0023003665],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011778245,0.00016519721,0.0073077474,0.00066589215,0.00020261959,0.0007060686,0.0010707818,0.30765662,0.010810012,0.3031952,0.009748748,0.35835335],"study_design_scores_gemma":[0.00002539573,0.00013684122,0.0027296122,0.0002674157,0.000088735804,0.00032668604,0.00027965385,0.8672855,0.0030227022,0.09006308,0.0356714,0.000102967635],"about_ca_topic_score_codex":0.016033804,"about_ca_topic_score_gemma":0.008374198,"teacher_disagreement_score":0.016033804,"about_ca_system_score_codex":0.0018833318,"about_ca_system_score_gemma":0.0024879635,"threshold_uncertainty_score":0.03389579},"labels":[],"label_agreement":null},{"id":"W1854452706","doi":"10.1007/978-3-642-03013-0_16","title":"Supporting Framework Use via Automatically Extracted Concept-Implementation Templates","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Template; Computer science; Documentation; Code (set theory); Software engineering; Proof of concept; Programming language; Set (abstract data type); Operating system","score_opus":0.02232149777068025,"score_gpt":0.3241798272673513,"score_spread":0.301858329496671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1854452706","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011612177,0.00025907735,0.92432153,0.00018965914,0.00010366964,0.0003922531,0.0016819463,0.055650957,0.0057887943],"genre_scores_gemma":[0.09868902,0.0003432653,0.88332343,0.00012379343,0.000031565596,0.00039512385,0.0045728264,0.009023325,0.0034976834],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99801743,0.00039387553,0.0002451434,0.0003616927,0.0008210112,0.00016093784],"domain_scores_gemma":[0.9886258,0.007524633,0.00070872507,0.0017910398,0.0011590221,0.00019077896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002989024,0.0016426945,0.000916082,0.0026323036,0.0005613431,0.0036411122,0.0026088979,0.0016714813,0.013284229],"category_scores_gemma":[0.024539175,0.0015942248,0.0016154996,0.0014565914,0.0005852902,0.004948815,0.0022232346,0.0021490883,0.005689054],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060754135,0.00047657595,0.0068555083,0.001932612,0.00018804445,0.0018167013,0.002655368,0.016292393,0.044771858,0.047758102,0.045457747,0.83118755],"study_design_scores_gemma":[0.0003030712,0.00019961061,0.0032934449,0.0015855682,0.0005404395,0.0027275148,0.00089211133,0.48993075,0.17941758,0.065361224,0.25542697,0.00032172812],"about_ca_topic_score_codex":0.0027398719,"about_ca_topic_score_gemma":0.005475537,"teacher_disagreement_score":0.013284229,"about_ca_system_score_codex":0.0006766771,"about_ca_system_score_gemma":0.0024147008,"threshold_uncertainty_score":0.04444021},"labels":[],"label_agreement":null},{"id":"W1857279498","doi":"10.1007/978-3-642-21043-3_49","title":"Intelligent Software Development Environments: Integrating Natural Language Processing with the Eclipse Platform","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software engineering; Software development; Eclipse; Software construction; Source code; Programming language; Commit; Software framework; Software; Software system; Static program analysis; Database","score_opus":0.01770433264318706,"score_gpt":0.23893936121384818,"score_spread":0.22123502857066113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1857279498","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010030557,0.0004152288,0.9613609,0.00019074939,0.0000565406,0.00007869065,0.00023907542,0.021424659,0.006203653],"genre_scores_gemma":[0.08395066,0.0008099852,0.90347517,0.0001677313,0.00003552631,0.0001038288,0.0019561716,0.003918163,0.005582788],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99931955,0.0001622778,0.00008540445,0.00013204181,0.00024038908,0.000060337145],"domain_scores_gemma":[0.9986663,0.00084280333,0.00009152938,0.00021852912,0.00012585806,0.000055040397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015365425,0.0005728176,0.00048728805,0.0007437632,0.00026945912,0.001816434,0.0015549437,0.00056994864,0.0019065475],"category_scores_gemma":[0.0029982931,0.0006121779,0.0008168544,0.00051161647,0.00040378398,0.002561262,0.0012231924,0.0013034436,0.0012518764],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005937033,0.00042286835,0.002242707,0.0007279989,0.00012641089,0.00060529384,0.0009940815,0.013026357,0.050678242,0.062516615,0.023266288,0.8447994],"study_design_scores_gemma":[0.00041422216,0.000391916,0.0040228623,0.000507145,0.0003649719,0.0024012777,0.00033913457,0.39058715,0.11764223,0.1300059,0.35313487,0.00018836245],"about_ca_topic_score_codex":0.0008922682,"about_ca_topic_score_gemma":0.0019478895,"teacher_disagreement_score":0.0019065475,"about_ca_system_score_codex":0.00030968204,"about_ca_system_score_gemma":0.00067926163,"threshold_uncertainty_score":0.00812614},"labels":[],"label_agreement":null},{"id":"W1861554568","doi":"10.1109/ms.2010.58","title":"The Theory of Relative Dependency: Higher Coupling Concentration in Smaller Modules","year":2010,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Code refactoring; Dependency (UML); Agile software development; Quality assurance; Quality (philosophy); Computer science; Scale (ratio); Reliability engineering; Software; Software engineering; Risk analysis (engineering); Engineering; Operations management; Business; Programming language","score_opus":0.017949501939768936,"score_gpt":0.25553881460230954,"score_spread":0.2375893126625406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1861554568","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4463314,0.002368939,0.47702843,0.005309214,0.00029421409,0.00035369452,0.0005893208,0.0007773984,0.06694731],"genre_scores_gemma":[0.96636915,0.0005609621,0.029280962,0.0010603446,0.00013185476,0.00024315769,0.00014264339,0.00010306977,0.0021078452],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9951115,0.0015227817,0.00033736124,0.001517402,0.0012206895,0.00029030527],"domain_scores_gemma":[0.9628974,0.023489838,0.006999013,0.0029619,0.002433845,0.0012180562],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004390896,0.00087813026,0.0006922959,0.0046149264,0.0011980075,0.0020396905,0.0016065395,0.0017492148,0.009050162],"category_scores_gemma":[0.03410671,0.00089553243,0.0012900241,0.0024782196,0.0054971417,0.005839703,0.0031659384,0.0015823895,0.0012160016],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000993435,0.00067103945,0.30776715,0.0016013308,0.00094634347,0.0017520869,0.011286009,0.0116967205,0.025931107,0.4493002,0.0063599753,0.18169461],"study_design_scores_gemma":[0.00026368306,0.0012944447,0.44338542,0.00031263067,0.0007297693,0.007607954,0.002599809,0.033918496,0.011488262,0.47942838,0.018695585,0.00027559136],"about_ca_topic_score_codex":0.0016375077,"about_ca_topic_score_gemma":0.0007503344,"teacher_disagreement_score":0.009050162,"about_ca_system_score_codex":0.0017527079,"about_ca_system_score_gemma":0.00094371685,"threshold_uncertainty_score":0.030275822},"labels":[],"label_agreement":null},{"id":"W1865143696","doi":"10.1109/icsm.1990.131382","title":"A model for estimating perfective software maintenance projects","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Function point; Function (biology); Point (geometry); Computer science; Software; Adaptation (eye); Mathematics; Software development; Programming language","score_opus":0.05777412581591333,"score_gpt":0.28172857928978284,"score_spread":0.2239544534738695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1865143696","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032207787,0.00012645649,0.96388,0.00019124229,0.000022477605,0.00011080717,0.00048952893,0.00043349553,0.0025382151],"genre_scores_gemma":[0.68028057,0.0006519351,0.3054989,0.00007424419,0.000067633446,0.0010364726,0.0016676531,0.00022579962,0.010496698],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975055,0.00086531095,0.00015659175,0.00058178947,0.000650542,0.00024020775],"domain_scores_gemma":[0.99245256,0.0048878435,0.0010922699,0.0004199049,0.0009884291,0.00015898928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004744698,0.0012738343,0.00080393115,0.002408819,0.0004923825,0.0025713926,0.0023689643,0.0016630881,0.003674701],"category_scores_gemma":[0.015907556,0.0011079219,0.0010864115,0.0015163579,0.0009810429,0.0026683868,0.001141998,0.0014732933,0.001105116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000042245832,0.00003876444,0.0020044036,0.000053965054,0.000038964026,0.000061904415,0.00011369134,0.9609622,0.00043511082,0.019746875,0.00057812326,0.015923763],"study_design_scores_gemma":[0.000011911105,0.000058588997,0.000829483,0.000013119521,0.000021381733,0.00005172892,0.000022084208,0.9860168,0.00019905853,0.0115938205,0.0011586156,0.000023284276],"about_ca_topic_score_codex":0.010722206,"about_ca_topic_score_gemma":0.0073111653,"teacher_disagreement_score":0.010722206,"about_ca_system_score_codex":0.0020037375,"about_ca_system_score_gemma":0.0016404517,"threshold_uncertainty_score":0.025092661},"labels":[],"label_agreement":null},{"id":"W1867669781","doi":"10.1109/case.1992.200166","title":"Implementing CASE tools: the need for a normative framework","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Normative; Computer science; Institution; Software engineering; Software; Knowledge management; Management science; Engineering; Programming language; Sociology","score_opus":0.04365527617068539,"score_gpt":0.3251922562653341,"score_spread":0.2815369800946487,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1867669781","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065911426,0.0013758872,0.85849077,0.050849535,0.00074138126,0.0010440319,0.0000572836,0.0007672981,0.08008258],"genre_scores_gemma":[0.1438514,0.0018466229,0.83557093,0.005824308,0.00056517735,0.0031905521,0.00030470267,0.00048101568,0.0083652595],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.7665385,0.15278968,0.01710661,0.007061352,0.0520349,0.004468924],"domain_scores_gemma":[0.69285655,0.18096666,0.013078614,0.034921024,0.069989935,0.008187266],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.24027738,0.0017731375,0.0014736139,0.009940247,0.011583748,0.042874344,0.008846987,0.013223041,0.0027250226],"category_scores_gemma":[0.22768556,0.0020943028,0.0011419142,0.004273752,0.048130717,0.05858308,0.014412189,0.014724878,0.0015515441],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000048874954,0.000045382472,0.00018329997,0.000071197835,0.0000045640654,0.000065849184,0.002324479,0.00073252857,0.00011960042,0.9805197,0.0015731303,0.014355432],"study_design_scores_gemma":[0.000019979892,0.000039777377,0.00015779243,0.00094781886,0.000009249508,0.00021710205,0.0038245106,0.004698859,0.0006312793,0.9080074,0.081384376,0.000061996994],"about_ca_topic_score_codex":0.0054154457,"about_ca_topic_score_gemma":0.0062497617,"teacher_disagreement_score":0.24027738,"about_ca_system_score_codex":0.012804735,"about_ca_system_score_gemma":0.035109606,"threshold_uncertainty_score":0.93687326},"labels":[],"label_agreement":null},{"id":"W1869933348","doi":"10.24908/pceea.v0i0.3164","title":"Plagiarism Detection in Code-Based Assignments","year":2010,"lang":"en","type":"article","venue":"Proceedings of the Canadian Engineering Education Association (CEEA)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Plagiarism detection; Computer science; Programming language; Code (set theory); Natural language processing","score_opus":0.0050967840772512445,"score_gpt":0.21209383837784127,"score_spread":0.20699705430059,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1869933348","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88940877,0.0019772488,0.08582918,0.00034161727,0.00020913934,0.0007670564,0.0027406232,0.0076061916,0.0111201955],"genre_scores_gemma":[0.9315219,0.0003279504,0.059406362,0.000088152825,0.00011741213,0.00027286255,0.0025298141,0.0003867514,0.0053488454],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9845345,0.0026763554,0.0014812927,0.002463757,0.00820759,0.0006364663],"domain_scores_gemma":[0.910009,0.032414492,0.023406215,0.009222754,0.021742703,0.0032048095],"candidate_categories":["metaresearch","research_integrity"],"consensus_categories":[],"category_scores_codex":[0.0051350314,0.00064516853,0.000988896,0.017345779,0.0014817945,0.002215637,0.0018143066,0.0015259942,0.0021913657],"category_scores_gemma":[0.06612677,0.00041567168,0.00035849697,0.0077574123,0.00084722624,0.0028182839,0.0023229336,0.0010286958,0.0024129075],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015736847,0.00070469565,0.281908,0.0011338054,0.00017163261,0.0014066909,0.0054961946,0.0043390947,0.05468344,0.0027061135,0.014876832,0.63099974],"study_design_scores_gemma":[0.0001443366,0.001298581,0.7015165,0.00041805528,0.00017835926,0.0062930165,0.0023221993,0.1479887,0.106077425,0.00777714,0.025610566,0.0003751187],"about_ca_topic_score_codex":0.0026876926,"about_ca_topic_score_gemma":0.0035121506,"teacher_disagreement_score":0.998474,"about_ca_system_score_codex":0.0011487757,"about_ca_system_score_gemma":0.0009741556,"threshold_uncertainty_score":0.027156949},"labels":[{"model":"gemma","categories":["research_integrity"],"domain":null,"study_design":"not_applicable","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":["research_integrity"],"domain":null,"study_design":"design_other","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W1870700426","doi":"10.1007/978-3-642-01680-6_12","title":"Evidence-Based Insights about Issue Management Processes: An Exploratory Study","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Context (archaeology); Process (computing); Data science; Exploratory research; Software development; Software development process; Empirical research; Software; Software engineering; Knowledge management","score_opus":0.042271127125060655,"score_gpt":0.2900931711198417,"score_spread":0.24782204399478103,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1870700426","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97074294,0.0010938691,0.011305502,0.0025366507,0.000023608538,0.00059101556,0.00016841115,0.000019340232,0.0135187125],"genre_scores_gemma":[0.9923327,0.0006243176,0.005937336,0.0001772758,0.000014497066,0.00022903964,0.000103998325,0.000011622111,0.00056913006],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.96818495,0.022401178,0.0023137135,0.0009613088,0.0054849,0.00065385824],"domain_scores_gemma":[0.38952258,0.5775694,0.017064005,0.0059275567,0.008390043,0.0015263491],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04383738,0.00038155535,0.0005462775,0.0039939317,0.0027906215,0.0065593356,0.0020705587,0.0020571228,0.0043488787],"category_scores_gemma":[0.2788437,0.00071744085,0.00039457748,0.004392382,0.0028814385,0.009128527,0.0039611054,0.0033453526,0.0003141111],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014154263,0.0030071123,0.1141615,0.0026194577,0.00015963281,0.002884901,0.67035204,0.0007071841,0.0032119912,0.03843159,0.0013852209,0.1616639],"study_design_scores_gemma":[0.0003915042,0.0028305755,0.15100804,0.0036439572,0.00032335005,0.003596805,0.74316674,0.0051513007,0.0048830034,0.0541114,0.030751897,0.0001414582],"about_ca_topic_score_codex":0.0007782161,"about_ca_topic_score_gemma":0.001238113,"teacher_disagreement_score":0.04383738,"about_ca_system_score_codex":0.0019165494,"about_ca_system_score_gemma":0.0042132344,"threshold_uncertainty_score":0.23183703},"labels":[],"label_agreement":null},{"id":"W1870786460","doi":"10.1109/ase.1998.732687","title":"Developing the designer's toolkit with software comprehension models","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Comprehension; Software; Task (project management); Program comprehension; Cognition; Software engineering; Human–computer interaction; Software system; Systems engineering; Programming language; Engineering","score_opus":0.0751801633652289,"score_gpt":0.24422481654289455,"score_spread":0.16904465317766565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1870786460","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003283127,0.000053949952,0.9913118,0.0005162781,0.000013428975,0.0000882942,0.000036928675,0.0017064448,0.0029897885],"genre_scores_gemma":[0.048322268,0.00014915079,0.9484998,0.00009203489,0.000013524441,0.00026564318,0.000110998015,0.0005205129,0.002026036],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.993471,0.0036603177,0.00053942547,0.0008889662,0.0012160292,0.00022419982],"domain_scores_gemma":[0.97958595,0.01320648,0.0007842835,0.004414131,0.0015646907,0.0004444218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011393923,0.0014982768,0.0011740593,0.0029745377,0.001047981,0.0063401023,0.004339274,0.0030181245,0.006116514],"category_scores_gemma":[0.038193807,0.002768945,0.0026563583,0.0010907473,0.003952793,0.015198838,0.0063973647,0.004878375,0.0029843212],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011754374,0.00033881972,0.0027206256,0.00080506684,0.0001702713,0.0005946287,0.013362532,0.05768998,0.0069877175,0.75176084,0.0058953725,0.15955658],"study_design_scores_gemma":[0.000122047524,0.00016923965,0.00038142127,0.00049874594,0.00011991793,0.000512758,0.0011397358,0.4027569,0.008372105,0.51308846,0.07272059,0.00011809143],"about_ca_topic_score_codex":0.0015729938,"about_ca_topic_score_gemma":0.002551922,"teacher_disagreement_score":0.011393923,"about_ca_system_score_codex":0.0015158765,"about_ca_system_score_gemma":0.0032696957,"threshold_uncertainty_score":0.060257554},"labels":[],"label_agreement":null},{"id":"W1872266072","doi":"10.1109/wpc.2005.11","title":"Browsing Software Architectures With LSEdit","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Visualization; Software visualization; Software analytics; Program comprehension; Software engineering; Graph; Software system; World Wide Web; Database; Programming language; Operating system; Software construction; Data mining; Theoretical computer science","score_opus":0.009422595660124485,"score_gpt":0.23332225491990768,"score_spread":0.2238996592597832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1872266072","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016414382,0.0011202269,0.6943788,0.0014181429,0.0002920853,0.00025341273,0.028538842,0.23392591,0.023658251],"genre_scores_gemma":[0.117044084,0.0021763335,0.80216813,0.00068074535,0.00014093088,0.000597239,0.04243176,0.019466883,0.015293843],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99973756,0.00005182485,0.000027261482,0.00004824562,0.00010960769,0.00002542465],"domain_scores_gemma":[0.9980355,0.0011884248,0.00010615088,0.0002679297,0.00027138056,0.00013062598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007639527,0.0012357221,0.0005516031,0.005837871,0.00048927194,0.002853386,0.0011594915,0.0010520201,0.028043697],"category_scores_gemma":[0.003702096,0.00060186605,0.00075744884,0.0028566306,0.0003760149,0.0027166565,0.0024312378,0.0014987206,0.0061258227],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006414725,0.00019050519,0.0030868463,0.0026801755,0.00015428051,0.0017829646,0.0033206877,0.011356784,0.02484207,0.03963181,0.33931345,0.572999],"study_design_scores_gemma":[0.00026640747,0.00013239327,0.0034784214,0.0003940962,0.00008783301,0.0018888866,0.0009000132,0.13900638,0.03276595,0.067497514,0.75342315,0.00015894602],"about_ca_topic_score_codex":0.0019452114,"about_ca_topic_score_gemma":0.004431406,"teacher_disagreement_score":0.028043697,"about_ca_system_score_codex":0.000591396,"about_ca_system_score_gemma":0.00068157987,"threshold_uncertainty_score":0.093815506},"labels":[],"label_agreement":null},{"id":"W1872675725","doi":"10.1109/icsm.2000.882972","title":"Bridging program comprehension tools by design navigation","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Program comprehension; Computer science; Suite; Bridging (networking); Software engineering; Human–computer interaction; Visualization; Perspective (graphical); Comprehension; Software; World Wide Web; Software system; Programming language; Artificial intelligence","score_opus":0.033685985196059975,"score_gpt":0.28903196195665043,"score_spread":0.25534597676059045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1872675725","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013640044,0.0006085754,0.96983117,0.0009129271,0.00006746689,0.0001872142,0.000058742047,0.007931433,0.006762413],"genre_scores_gemma":[0.06876243,0.00063099177,0.9241014,0.00043693706,0.000045134442,0.00033926006,0.00022634813,0.0020777958,0.0033797594],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98574394,0.008169898,0.0013822868,0.0015784629,0.0025240348,0.000601416],"domain_scores_gemma":[0.937263,0.04174388,0.0029766515,0.01399982,0.003251441,0.00076515804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014933592,0.0022683374,0.0012404752,0.0056817355,0.0012371907,0.0064920327,0.0038408139,0.004743117,0.007625928],"category_scores_gemma":[0.056898564,0.0021018032,0.0012111841,0.002845444,0.004000335,0.013286103,0.011449875,0.0041937325,0.0034642646],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031435254,0.0008287553,0.0035355017,0.0014083158,0.000063705214,0.0015115985,0.03188511,0.0052599255,0.026274621,0.13082927,0.009029138,0.7890597],"study_design_scores_gemma":[0.00045816676,0.0011751439,0.0036826024,0.0025662533,0.0002861649,0.0076898555,0.0077996235,0.109150186,0.06201663,0.29358715,0.51094586,0.0006424101],"about_ca_topic_score_codex":0.0011191906,"about_ca_topic_score_gemma":0.0012933827,"teacher_disagreement_score":0.014933592,"about_ca_system_score_codex":0.0009383364,"about_ca_system_score_gemma":0.002858095,"threshold_uncertainty_score":0.07897729},"labels":[],"label_agreement":null},{"id":"W1874892428","doi":"10.1007/3-540-45140-4_2","title":"Why Is It So Difficult to Introduce RE Research Results into Mainstream RE Practice?","year":2000,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada Research Chairs; University of Toronto; University of New Brunswick","funders":"","keywords":"Mainstream; Computer science; Data science","score_opus":0.03645989557913309,"score_gpt":0.3316178737241077,"score_spread":0.29515797814497463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1874892428","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029169858,0.033005726,0.25164193,0.50502956,0.010873868,0.00016035695,0.00009005382,0.0019411829,0.16808744],"genre_scores_gemma":[0.64199305,0.025971498,0.16557446,0.0798701,0.011449171,0.00050159,0.0001798053,0.0020233481,0.072436914],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.96649015,0.019068155,0.0015124358,0.0028207586,0.008533423,0.0015750853],"domain_scores_gemma":[0.8050731,0.14433828,0.0049398458,0.013971831,0.027799156,0.0038777662],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04953112,0.00085608294,0.0011763055,0.0032605908,0.0031624655,0.017798003,0.0029321343,0.008812796,0.006606443],"category_scores_gemma":[0.09513797,0.0007992416,0.0005832787,0.0026500167,0.015806548,0.051783808,0.007374431,0.011342999,0.0040954067],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007327018,0.00021873225,0.0012760736,0.0013228204,0.0000515826,0.0003244026,0.015846372,0.00051935413,0.0021857691,0.76910216,0.04088784,0.16819164],"study_design_scores_gemma":[0.00005679543,0.00014429976,0.0012060323,0.0018415956,0.00003616874,0.0007439027,0.020051718,0.003002665,0.0031938872,0.60360664,0.3660436,0.00007267398],"about_ca_topic_score_codex":0.0013158695,"about_ca_topic_score_gemma":0.002089161,"teacher_disagreement_score":0.9504689,"about_ca_system_score_codex":0.0030825268,"about_ca_system_score_gemma":0.003563376,"threshold_uncertainty_score":0.2619487},"labels":[],"label_agreement":null},{"id":"W1879993036","doi":"10.1109/sess.1999.766599","title":"A structured analysis of the new ISO standard on functional size measurement-definition of concepts","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Standardization; Strengths and weaknesses; Computer science; Completeness (order theory); Process (computing); Software; Software engineering; Set (abstract data type); Software measurement; Perspective (graphical); Contrast (vision); Systems engineering; Data mining; Software quality; Software development; Engineering; Mathematics; Artificial intelligence; Programming language","score_opus":0.041974586945583515,"score_gpt":0.26931814227001644,"score_spread":0.22734355532443293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1879993036","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019027852,0.0032630607,0.91809183,0.0060003446,0.0005502213,0.0011308519,0.00088225224,0.000415825,0.050637744],"genre_scores_gemma":[0.098578066,0.003970631,0.88358617,0.0014446571,0.00038985847,0.0019572212,0.0023789948,0.0003117323,0.0073826415],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96939385,0.007299483,0.0022862544,0.0016728916,0.01873514,0.00061235495],"domain_scores_gemma":[0.96859556,0.011015393,0.0022949337,0.002967938,0.014808229,0.000317963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019041287,0.0010815787,0.0008612227,0.010959143,0.0016885196,0.0070637087,0.0015267528,0.0016092192,0.0022418352],"category_scores_gemma":[0.033337414,0.00074977614,0.0014467983,0.007622476,0.0047156448,0.009611377,0.0030842551,0.0034242151,0.0010154509],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000056993777,0.00011065023,0.0019327094,0.0009561904,0.00002813312,0.00024778297,0.0056774886,0.004359937,0.006445877,0.78353536,0.010223184,0.18642564],"study_design_scores_gemma":[0.00004200266,0.00039463662,0.0092671625,0.0035352588,0.000089633446,0.0011365935,0.007615687,0.041230954,0.012056086,0.49442697,0.43000546,0.00019947572],"about_ca_topic_score_codex":0.0031287456,"about_ca_topic_score_gemma":0.0030334627,"teacher_disagreement_score":0.019041287,"about_ca_system_score_codex":0.005449267,"about_ca_system_score_gemma":0.012784313,"threshold_uncertainty_score":0.10070115},"labels":[],"label_agreement":null},{"id":"W1880500695","doi":"10.1002/spe.2122","title":"Fast and effective soft links","year":2012,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Hypertext; Markup language; Source code; Java; Code (set theory); Traceability; Information retrieval; Signature (topology); Programming language; Database; World Wide Web; XML; Software engineering","score_opus":0.012176318863095538,"score_gpt":0.29793788636462176,"score_spread":0.2857615675015262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1880500695","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021604855,0.00028146736,0.9577361,0.00039284528,0.00012318487,0.00034157757,0.00042902876,0.012541592,0.0065494296],"genre_scores_gemma":[0.22590671,0.0002895645,0.7528395,0.00037688971,0.00017647115,0.00048552227,0.0016754671,0.0031917044,0.015058202],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9900287,0.002714394,0.00081289513,0.0015205503,0.0045000305,0.00042344144],"domain_scores_gemma":[0.96472245,0.016003486,0.0021975287,0.010697063,0.005475499,0.00090405333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042836643,0.0012114587,0.0010941724,0.004620007,0.0018006013,0.0054204124,0.003013221,0.0019139899,0.013729931],"category_scores_gemma":[0.038811993,0.0006937886,0.0009807275,0.0032900476,0.0021151104,0.00751621,0.0062857396,0.002385772,0.006406049],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053040107,0.00036635302,0.0028213738,0.0008021154,0.000112120426,0.00073660247,0.001877088,0.018368268,0.026015941,0.05566914,0.019506492,0.8731941],"study_design_scores_gemma":[0.00032300683,0.00068336277,0.0033170567,0.00048279582,0.00021108973,0.0015520677,0.0019028938,0.38337958,0.19572785,0.23827416,0.1737975,0.00034864666],"about_ca_topic_score_codex":0.0013961636,"about_ca_topic_score_gemma":0.0014784188,"teacher_disagreement_score":0.013729931,"about_ca_system_score_codex":0.00081895676,"about_ca_system_score_gemma":0.0015933067,"threshold_uncertainty_score":0.04593122},"labels":[],"label_agreement":null},{"id":"W1882295992","doi":"10.1109/nafips.2001.944298","title":"Evaluating software project similarity by using linguistic quantifier guided aggregations","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Quantifier (linguistics); Software; Similarity (geometry); Set (abstract data type); Fuzzy logic; Natural language processing; Fuzzy set; Artificial intelligence; Programming language; Linguistics; Software engineering","score_opus":0.22279365020496067,"score_gpt":0.40811152239580223,"score_spread":0.18531787219084156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1882295992","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38778934,0.0003852412,0.6083004,0.00024751294,0.00004175722,0.00021705247,0.0002802203,0.00058923155,0.002149205],"genre_scores_gemma":[0.7442778,0.000105061714,0.25477633,0.000024648216,0.000039686023,0.00013088403,0.00030003858,0.000029636503,0.0003159663],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916374,0.0026026722,0.0009971398,0.00063882896,0.0038966862,0.00022737973],"domain_scores_gemma":[0.9788094,0.010452482,0.004120074,0.001169152,0.0048380666,0.00061086495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066892733,0.00062306516,0.00092279894,0.00909214,0.00073259586,0.0020765292,0.0008768933,0.00056791946,0.000703557],"category_scores_gemma":[0.029611452,0.00024441755,0.00059947185,0.005070656,0.00067143264,0.0029090093,0.0014751133,0.00059864594,0.000114537426],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011796678,0.00060371397,0.06221173,0.00064642366,0.00069323735,0.00041179277,0.002461414,0.12968817,0.026091304,0.053654015,0.0026733296,0.71968514],"study_design_scores_gemma":[0.000092186245,0.00037308165,0.024072783,0.00006220269,0.00018026862,0.00018791936,0.00063660514,0.9161815,0.01528735,0.040731408,0.0020727243,0.00012191817],"about_ca_topic_score_codex":0.0028150026,"about_ca_topic_score_gemma":0.0032688612,"teacher_disagreement_score":0.00909214,"about_ca_system_score_codex":0.0012010022,"about_ca_system_score_gemma":0.0011514259,"threshold_uncertainty_score":0.035376728},"labels":[],"label_agreement":null},{"id":"W188835152","doi":"","title":"A formalism of ontology to support a software maintenance knowledge based system","year":2006,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software engineering; Computer science; Software maintenance; Software mining; Capability Maturity Model; Software development; Ontology; Software system; Software construction; Team software process; Software; Knowledge management; Operating system","score_opus":0.00873206749947429,"score_gpt":0.22707239515293123,"score_spread":0.21834032765345696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W188835152","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019195747,0.00021201505,0.987959,0.0011269267,0.00012079286,0.00019399419,0.00033634176,0.00079184654,0.0073394575],"genre_scores_gemma":[0.03342323,0.0003603283,0.96256006,0.00033005315,0.000059009813,0.0003724507,0.000653086,0.00014835349,0.002093393],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99650025,0.0011501658,0.00073593896,0.0004980776,0.0008440769,0.00027160704],"domain_scores_gemma":[0.9949421,0.0020604797,0.0004912075,0.0013247415,0.0007804947,0.00040088227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008508374,0.0007442669,0.0011522878,0.0046584653,0.0034601667,0.0071347086,0.003167385,0.0025195386,0.003734989],"category_scores_gemma":[0.009957799,0.001183185,0.0035402605,0.0053745615,0.004977178,0.013419617,0.005718189,0.005532286,0.0012640475],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000016112139,0.000041587627,0.0003433009,0.00011581042,0.000034898338,0.00027927847,0.0020431455,0.0030017402,0.0010546631,0.9671048,0.0030009402,0.022963751],"study_design_scores_gemma":[0.00006198584,0.00005637011,0.00046008316,0.00037305782,0.00013780582,0.0006765074,0.000910438,0.049475025,0.0020488324,0.6937903,0.2519143,0.00009523245],"about_ca_topic_score_codex":0.015760718,"about_ca_topic_score_gemma":0.014291872,"teacher_disagreement_score":0.015760718,"about_ca_system_score_codex":0.0033198206,"about_ca_system_score_gemma":0.0073107984,"threshold_uncertainty_score":0.044997096},"labels":[],"label_agreement":null},{"id":"W1891539143","doi":"","title":"Task-directed software inspection technique: an experiment and case study","year":2000,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Software inspection; Task (project management); Software engineering; Computer science; Software; Software construction; Code (set theory); Sample (material); Software development; Engineering; Software quality; Systems engineering; Programming language","score_opus":0.08221613540789634,"score_gpt":0.40925681133602393,"score_spread":0.3270406759281276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1891539143","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98256487,0.00009941497,0.014212656,0.00010411269,0.000024749479,0.0016403588,0.00010568316,0.00013704148,0.001111209],"genre_scores_gemma":[0.9147339,0.00029573252,0.07967026,0.00019097005,0.000028413599,0.0027425804,0.00021135273,0.00006388891,0.0020628623],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9942001,0.0029071646,0.0005381584,0.0008991948,0.0010536611,0.00040168522],"domain_scores_gemma":[0.94601506,0.043494415,0.0019214336,0.0041413596,0.0034741145,0.0009535461],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007942671,0.00092106365,0.0009260438,0.0007985408,0.0009379417,0.0007588179,0.0017097847,0.002794321,0.0017558],"category_scores_gemma":[0.02426695,0.0005932299,0.0006970566,0.0006613526,0.0013473289,0.0013711955,0.0010920407,0.0015329433,0.00046205492],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012786055,0.16808696,0.04416426,0.004554376,0.0003880747,0.00678258,0.09403415,0.01674493,0.3305026,0.0053253775,0.0052411803,0.31138942],"study_design_scores_gemma":[0.009701755,0.32564813,0.10922521,0.0008407611,0.00090143055,0.01003117,0.032459144,0.11662286,0.353879,0.008216805,0.03147981,0.0009939244],"about_ca_topic_score_codex":0.0012685487,"about_ca_topic_score_gemma":0.0018762605,"teacher_disagreement_score":0.007942671,"about_ca_system_score_codex":0.0005414653,"about_ca_system_score_gemma":0.0009318756,"threshold_uncertainty_score":0.04200542},"labels":[],"label_agreement":null},{"id":"W1893729795","doi":"10.1007/978-3-540-72530-5_9","title":"A Rough-Hybrid Approach to Software Defect Classification","year":2007,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Winnipeg","funders":"","keywords":"Computer science; Rough set; Data mining; Decision tree; Software; Fuzzy logic; Artificial intelligence; Dominance-based rough set approach; Knowledge extraction; Machine learning; Fuzzy set","score_opus":0.04320667514004044,"score_gpt":0.2811904742937011,"score_spread":0.23798379915366064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1893729795","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007006812,0.00043634034,0.99115705,0.00011177296,0.000036089066,0.000042498334,0.000070789676,0.00027453006,0.0008642089],"genre_scores_gemma":[0.15149903,0.00049238506,0.84327364,0.00013694468,0.0001572767,0.0001791182,0.00039263663,0.0001026928,0.0037662557],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960598,0.0010957257,0.00025948862,0.0003693989,0.002063355,0.00015229329],"domain_scores_gemma":[0.9944992,0.003296628,0.00021845673,0.0006391926,0.0012380253,0.00010841012],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031049203,0.000832985,0.002062406,0.004081918,0.0007590587,0.0023801292,0.0026596298,0.0012016388,0.002092462],"category_scores_gemma":[0.006532843,0.0006850691,0.00191837,0.0030902426,0.0009703125,0.002328984,0.0015446983,0.0016062366,0.00069337146],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002745056,0.0004048211,0.0028654176,0.0004606478,0.0005424182,0.00026088764,0.00029557158,0.20389618,0.009746861,0.053517148,0.0071767615,0.7205587],"study_design_scores_gemma":[0.000019301864,0.00012135402,0.0012953883,0.000035516863,0.00011301638,0.00019812779,0.00007116395,0.9428521,0.0019812977,0.050511584,0.0027562566,0.000044894397],"about_ca_topic_score_codex":0.0023598694,"about_ca_topic_score_gemma":0.0032835025,"teacher_disagreement_score":0.004081918,"about_ca_system_score_codex":0.0007876768,"about_ca_system_score_gemma":0.00084034586,"threshold_uncertainty_score":0.016420603},"labels":[],"label_agreement":null},{"id":"W1894973801","doi":"10.1007/978-3-540-75975-1_9","title":"Recovering Business Rules from Legacy Source Code for System Modernization","year":2007,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Legacy system; Semantics of Business Vocabulary and Business Rules; Business rule; Source code; Computer science; Modernization theory; Software engineering; Documentation; Legacy code; Reverse engineering; Code (set theory); Business logic; KPI-driven code analysis; Business process; Programming language; Engineering; Static program analysis; Software development; Set (abstract data type); Operations management; Software; Work in process; Political science; Law","score_opus":0.029155691322831374,"score_gpt":0.26073486195172746,"score_spread":0.23157917062889607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1894973801","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15585981,0.0012582146,0.7735848,0.0013282243,0.0002913259,0.00030929016,0.002145929,0.05646004,0.008762381],"genre_scores_gemma":[0.3937428,0.00092430494,0.5872751,0.00026516238,0.00008930739,0.00011390497,0.0052541452,0.0053395815,0.0069957124],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980659,0.0002525064,0.000114967625,0.00030300365,0.0010835471,0.00018013854],"domain_scores_gemma":[0.9898699,0.0027817278,0.00078591186,0.0048391577,0.0016064927,0.00011684691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013923985,0.0012472229,0.00084008207,0.0030134604,0.0006021805,0.0019916801,0.0018782966,0.0012188596,0.0031730898],"category_scores_gemma":[0.01362543,0.0012187196,0.0015397246,0.002007352,0.00083198113,0.003179821,0.001967872,0.0027789369,0.0023420637],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027947777,0.00031207685,0.010082047,0.0004264865,0.0001956413,0.000847577,0.00042545708,0.05331551,0.03991566,0.0101296725,0.017170483,0.8668998],"study_design_scores_gemma":[0.000062271574,0.00013230616,0.0057972935,0.00013431984,0.0002505415,0.00074329757,0.0002848669,0.80412465,0.1213982,0.04098166,0.026002074,0.0000885837],"about_ca_topic_score_codex":0.002693932,"about_ca_topic_score_gemma":0.004806606,"teacher_disagreement_score":0.0031730898,"about_ca_system_score_codex":0.0007214239,"about_ca_system_score_gemma":0.0019114099,"threshold_uncertainty_score":0.010614991},"labels":[],"label_agreement":null},{"id":"W1895441878","doi":"10.4230/dagsemproc.06301.7","title":"Detection of Plagiarism in University Projects Using Metrics-based Spectral Similarity","year":2007,"lang":"en","type":"article","venue":"DROPS (Schloss Dagstuhl – Leibniz Center for Informatics)","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Plagiarism detection; Similarity (geometry); Computer science; Spectral analysis; Source code; Data mining; Clan; Information retrieval; Artificial intelligence; Programming language; Physics","score_opus":0.02642885969981242,"score_gpt":0.2678502540887758,"score_spread":0.24142139438896335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1895441878","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5270267,0.0008200661,0.46379298,0.00020459338,0.00006234107,0.0002949605,0.0007783643,0.0040552183,0.0029647653],"genre_scores_gemma":[0.8299459,0.000194052,0.1675661,0.00002430027,0.000040336614,0.00017815502,0.0011107157,0.00016644067,0.000773883],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9939207,0.0009898143,0.0005598473,0.0008513083,0.0034076502,0.0002707743],"domain_scores_gemma":[0.9833477,0.0049860324,0.0034326422,0.0019278723,0.0055514807,0.0007542248],"candidate_categories":["research_integrity"],"consensus_categories":[],"category_scores_codex":[0.0025432282,0.00057959347,0.001066194,0.014273492,0.000858112,0.0017462083,0.0009041701,0.0008156213,0.0010187237],"category_scores_gemma":[0.020652397,0.00026733775,0.00054036995,0.008216036,0.0007084139,0.0024181895,0.0016114409,0.00063024973,0.00073709304],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009378669,0.00040368555,0.09253872,0.0005399409,0.00028456835,0.00045471615,0.0016342521,0.013780389,0.09432578,0.0061482885,0.0032864208,0.78566533],"study_design_scores_gemma":[0.00011113052,0.0010977943,0.20393214,0.00011177677,0.00015985862,0.0049328757,0.0014655268,0.64205885,0.11590693,0.019033268,0.010872357,0.0003175608],"about_ca_topic_score_codex":0.0012220285,"about_ca_topic_score_gemma":0.0013265155,"teacher_disagreement_score":0.99918437,"about_ca_system_score_codex":0.00077122776,"about_ca_system_score_gemma":0.000704962,"threshold_uncertainty_score":0.0134500265},"labels":[],"label_agreement":null},{"id":"W1896715963","doi":"10.1109/iwrsp.1990.144033","title":"001: a rapid development approach for rapid prototyping based on a system that supports its own life cycle","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hamilton Health Sciences","funders":"","keywords":"Computer science; Obsolescence; Traceability; Reusability; Flexibility (engineering); Software engineering; Development (topology); Process (computing); Systems development life cycle; Interface (matter); Control reconfiguration; Programming language; Software development process; Embedded system; Software development; Operating system","score_opus":0.04786903910550738,"score_gpt":0.24251985447286184,"score_spread":0.19465081536735446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1896715963","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003733795,0.00018079876,0.9314077,0.00049615314,0.00024203911,0.0008269271,0.00009752711,0.008449563,0.05456561],"genre_scores_gemma":[0.028417751,0.00029284856,0.93838066,0.00026586367,0.000048110734,0.00091721915,0.00020244773,0.0017447759,0.029730307],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958488,0.001781706,0.00015909002,0.00038589202,0.001576697,0.000247833],"domain_scores_gemma":[0.99661785,0.0014369736,0.00019023275,0.00091342727,0.0005771824,0.0002643708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004762203,0.0012270316,0.0004460961,0.0012923374,0.00085722696,0.004096846,0.002722664,0.0018584976,0.015565663],"category_scores_gemma":[0.009784347,0.00101438,0.0008652424,0.00047592257,0.0020633521,0.0036236288,0.0036139146,0.0021854255,0.0064667854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040728977,0.00030727912,0.0009851663,0.0011039712,0.000068581496,0.0015383795,0.0058649895,0.012635939,0.070137694,0.41483855,0.053053677,0.4390585],"study_design_scores_gemma":[0.00023983845,0.00056037016,0.000744843,0.00036373295,0.00006621713,0.002838468,0.0003826536,0.038132697,0.030760447,0.05827443,0.8674496,0.00018680312],"about_ca_topic_score_codex":0.0011194742,"about_ca_topic_score_gemma":0.001626727,"teacher_disagreement_score":0.015565663,"about_ca_system_score_codex":0.0010879023,"about_ca_system_score_gemma":0.0021870192,"threshold_uncertainty_score":0.052072346},"labels":[],"label_agreement":null},{"id":"W1897334947","doi":"10.1109/hicss.1991.184094","title":"Measuring the quality of user-developed applications","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal","funders":"","keywords":"Cronbach's alpha; Quality (philosophy); Reliability (semiconductor); Computer science; Consistency (knowledge bases); Software quality; Internal consistency; Software; Measure (data warehouse); Construct validity; Test (biology); Data mining; Psychometrics; Software development; Mathematics; Artificial intelligence; Statistics; Programming language","score_opus":0.1512391064111098,"score_gpt":0.3279987561874381,"score_spread":0.17675964977632833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1897334947","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98188066,0.0004865686,0.008881106,0.00034251425,0.00001533535,0.000114100054,0.00020800212,0.00020597424,0.007865684],"genre_scores_gemma":[0.9890003,0.00024443716,0.009096124,0.00006284549,0.0000124396165,0.0000480542,0.00030321095,0.000029690693,0.0012027604],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9906294,0.0030323386,0.000987001,0.0004519401,0.0045446227,0.0003546206],"domain_scores_gemma":[0.9094344,0.040021937,0.012948516,0.005554019,0.02968742,0.0023537648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007644738,0.00026583055,0.00028087912,0.0025118648,0.00035052848,0.0026333954,0.0005007662,0.00059496425,0.0014166407],"category_scores_gemma":[0.07506398,0.0002734702,0.00032608834,0.0018166872,0.00055266195,0.002227794,0.0009294513,0.0006002111,0.0004060721],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035612605,0.00065931474,0.72516656,0.0004568439,0.00021041579,0.0001679377,0.009372894,0.0013704924,0.009315277,0.0016195179,0.002512996,0.24879162],"study_design_scores_gemma":[0.000048037444,0.0014706895,0.9543281,0.0002440027,0.0001388147,0.0008617652,0.0043386756,0.010771367,0.012070873,0.0018069326,0.013813672,0.00010707798],"about_ca_topic_score_codex":0.001317562,"about_ca_topic_score_gemma":0.0014786349,"teacher_disagreement_score":0.007644738,"about_ca_system_score_codex":0.0007550177,"about_ca_system_score_gemma":0.0008817279,"threshold_uncertainty_score":0.04042977},"labels":[],"label_agreement":null},{"id":"W1899857608","doi":"10.1109/tse.2015.2431680","title":"Facilitating Coordination between Software Developers: A Study and Techniques for Timely and Efficient Recommendations","year":2015,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"National Science Foundation","keywords":"Computer science; Software engineering; Software; Schedule; Software development; Software project management; Process management; Software construction; Engineering","score_opus":0.041973657220534324,"score_gpt":0.2945908783976059,"score_spread":0.2526172211770716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1899857608","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27755097,0.00178326,0.70187277,0.003298127,0.00010742497,0.0010883299,0.0003327161,0.0058549773,0.008111357],"genre_scores_gemma":[0.5807394,0.0006462506,0.4153025,0.0001705887,0.00005480064,0.00043215998,0.0003556384,0.00035039664,0.0019482294],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9695877,0.017693575,0.0022742779,0.004335615,0.0052026035,0.00090621854],"domain_scores_gemma":[0.83027154,0.117902964,0.01625291,0.020363905,0.01285437,0.0023542906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020773826,0.0013716507,0.0010091388,0.0045854435,0.00260345,0.0037432506,0.0029282141,0.001874008,0.0018041072],"category_scores_gemma":[0.14197183,0.0014909059,0.000820611,0.0033724254,0.0015330906,0.008288046,0.0024542417,0.0025059995,0.00090850657],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008415842,0.0009429084,0.07088503,0.0011508883,0.00018691606,0.0008115358,0.034007676,0.019014461,0.015587656,0.010856344,0.007265748,0.8384491],"study_design_scores_gemma":[0.00076289166,0.0023980797,0.076923914,0.0015993863,0.0008332223,0.002440135,0.03516954,0.6829632,0.04931013,0.04523197,0.101715244,0.0006523175],"about_ca_topic_score_codex":0.012911481,"about_ca_topic_score_gemma":0.011544011,"teacher_disagreement_score":0.020773826,"about_ca_system_score_codex":0.0022586046,"about_ca_system_score_gemma":0.0051949862,"threshold_uncertainty_score":0.10986382},"labels":[],"label_agreement":null},{"id":"W1902003010","doi":"10.1002/smr.1696","title":"A selection of distinguished papers from the 19th Working Conference on Reverse Engineering 2012","year":2014,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Reverse engineering; Computer science; Program slicing; Variety (cybernetics); Slicing; Selection (genetic algorithm); Java; Software engineering; Source code; Quality (philosophy); Data science; World Wide Web; Programming language; Artificial intelligence","score_opus":0.01548924074106418,"score_gpt":0.24196808550595889,"score_spread":0.22647884476489472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1902003010","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032306886,0.12377137,0.046147685,0.06090917,0.6659358,0.0010505739,0.0043721325,0.0021448145,0.092437685],"genre_scores_gemma":[0.014173801,0.1192434,0.030108659,0.012651818,0.26988578,0.0013062692,0.018046837,0.005218727,0.52936465],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99239576,0.0008186241,0.0006356015,0.000998046,0.0044156956,0.0007362514],"domain_scores_gemma":[0.9669543,0.0044472697,0.0014061098,0.0016957906,0.019257544,0.006239063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008634128,0.0025507258,0.0025719015,0.01252733,0.0028381245,0.010415226,0.0022432029,0.0029326836,0.08803973],"category_scores_gemma":[0.019370208,0.0009049109,0.002352334,0.010296399,0.0010081913,0.0062193703,0.00474032,0.0047581997,0.054812193],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009611557,0.000055311037,0.00018788713,0.00055899704,0.000022517155,0.000093189956,0.00007562646,0.00022403685,0.001285423,0.0020531737,0.8646935,0.13065419],"study_design_scores_gemma":[0.00002359081,0.000071468494,0.0005642013,0.00049508933,0.000021649115,0.00016929606,0.0001031473,0.00025065284,0.0007235622,0.0022038475,0.99534106,0.00003254787],"about_ca_topic_score_codex":0.0012798638,"about_ca_topic_score_gemma":0.0024947266,"teacher_disagreement_score":0.08803973,"about_ca_system_score_codex":0.0036199791,"about_ca_system_score_gemma":0.0047145663,"threshold_uncertainty_score":0.29452223},"labels":[],"label_agreement":null},{"id":"W1909156290","doi":"10.1109/wpc.2005.28","title":"Presenting Micro-Theories of Program Comprehension in Pattern Form","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Program comprehension; Comprehension; Computer science; Cognition; Set (abstract data type); Focus (optics); Pattern language (formal languages); Key (lock); Field (mathematics); Cognitive science; Artificial intelligence; Programming language; Software; Psychology; Software system; Mathematics","score_opus":0.017367170888316755,"score_gpt":0.2907227528323307,"score_spread":0.27335558194401394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1909156290","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008289798,0.0012700335,0.9363057,0.004194834,0.00012404827,0.00018717478,0.00019698251,0.0005380025,0.048893444],"genre_scores_gemma":[0.38891026,0.0030462204,0.5807838,0.0017299729,0.0004409208,0.0018626518,0.00082587654,0.00039034797,0.022009995],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99760085,0.00085443107,0.00023901407,0.00039796243,0.0006846285,0.00022305858],"domain_scores_gemma":[0.9941047,0.003235672,0.00066081743,0.0011739611,0.00059655035,0.00022830836],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027815823,0.0011506442,0.0007664812,0.003664636,0.0019495655,0.0064281453,0.0025583461,0.0027942476,0.017497692],"category_scores_gemma":[0.010513677,0.0008543872,0.0021846301,0.0047588143,0.008616717,0.019873045,0.00413959,0.0041376464,0.0024472696],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000072450252,0.000016220854,0.00035112273,0.00007785735,0.000008627377,0.00010134892,0.0012608881,0.0012562675,0.00019722167,0.9862576,0.0011484355,0.009317234],"study_design_scores_gemma":[0.0000088147235,0.00002082955,0.0002022549,0.000067598216,0.000009582227,0.00018717958,0.00043442866,0.009878827,0.0003202175,0.97257835,0.016281066,0.0000108879285],"about_ca_topic_score_codex":0.002307034,"about_ca_topic_score_gemma":0.0023260345,"teacher_disagreement_score":0.017497692,"about_ca_system_score_codex":0.0025585066,"about_ca_system_score_gemma":0.0013868443,"threshold_uncertainty_score":0.058535576},"labels":[],"label_agreement":null},{"id":"W1909497710","doi":"10.1007/s10664-015-9396-2","title":"Towards building a universal defect prediction model with rank transformed predictors","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":116,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"University of Victoria","keywords":"Computer science; Software; Rank (graph theory); Workflow; Predictive modelling; Context (archaeology); Eclipse; Data mining; Software development; Obstacle; Software bug; Software engineering; Machine learning; Database; Programming language; Mathematics","score_opus":0.02760184776548904,"score_gpt":0.2626829102734268,"score_spread":0.23508106250793773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1909497710","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031494536,0.0002372748,0.9655889,0.00030163702,0.000042664495,0.000047555983,0.0002721072,0.001329428,0.00068586366],"genre_scores_gemma":[0.58978766,0.00060711795,0.40218198,0.00038332978,0.00020293785,0.00025798884,0.0017893796,0.00028574548,0.00450384],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99805593,0.000593817,0.00015052508,0.00058825343,0.00039446022,0.0002171509],"domain_scores_gemma":[0.99412954,0.002699067,0.0005349945,0.0010987791,0.001313796,0.00022391217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046249214,0.0011733214,0.0021668072,0.0018499715,0.00059355254,0.001742359,0.0023338355,0.0015373337,0.001880942],"category_scores_gemma":[0.01195783,0.000783088,0.0013522485,0.0018841971,0.00098097,0.0031924692,0.0027558056,0.0023464027,0.001605855],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029182996,0.0005802307,0.023281325,0.0002691939,0.0004952291,0.00038424198,0.00035554304,0.5150514,0.008271501,0.039042983,0.006487436,0.40548903],"study_design_scores_gemma":[0.000008381884,0.000050160972,0.0008098582,0.00002082588,0.000040794504,0.00004603464,0.000020708576,0.98357683,0.0006385483,0.014135969,0.00063833763,0.000013570464],"about_ca_topic_score_codex":0.005129556,"about_ca_topic_score_gemma":0.0060565057,"teacher_disagreement_score":0.005129556,"about_ca_system_score_codex":0.000569749,"about_ca_system_score_gemma":0.002172701,"threshold_uncertainty_score":0.024459183},"labels":[],"label_agreement":null},{"id":"W191168329","doi":"","title":"A multidimensional empirical study on refactoring activity","year":2013,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; Concordia University","funders":"","keywords":"Code refactoring; Computer science; Empirical research; Software engineering; Code (set theory); Programming language; Software; Mathematics; Statistics","score_opus":0.18390792647154403,"score_gpt":0.45032773050272135,"score_spread":0.2664198040311773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W191168329","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99332666,0.00026402224,0.003608901,0.00013763079,0.0000055592404,0.00008575857,0.0002455484,0.000018468312,0.002307365],"genre_scores_gemma":[0.99611807,0.00015181351,0.0027997533,0.000027485014,0.000008591688,0.00014667798,0.00039759497,0.000015282621,0.0003347287],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9776372,0.0127173755,0.0025947478,0.0017951657,0.004425659,0.00082992384],"domain_scores_gemma":[0.7099877,0.20963462,0.04454049,0.012063362,0.02012206,0.0036518408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016171344,0.00037074322,0.00039640567,0.0052986243,0.0010276977,0.0026569674,0.0008011005,0.0010004287,0.0019201693],"category_scores_gemma":[0.1007251,0.00038137333,0.00054862903,0.0065119956,0.0015418299,0.0034228114,0.0020903314,0.0012259551,0.0005226302],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015883554,0.0006032056,0.9392764,0.000284542,0.000102237886,0.00022107891,0.022297462,0.00053131115,0.0011331331,0.0012131484,0.00030271034,0.033875998],"study_design_scores_gemma":[0.000016274964,0.00042811263,0.9679108,0.0001510354,0.000029584004,0.00041723775,0.022974815,0.0022285099,0.00097179855,0.00092178246,0.0038941398,0.00005595515],"about_ca_topic_score_codex":0.0013339084,"about_ca_topic_score_gemma":0.0014066303,"teacher_disagreement_score":0.016171344,"about_ca_system_score_codex":0.0012716576,"about_ca_system_score_gemma":0.0010235289,"threshold_uncertainty_score":0.08552331},"labels":[],"label_agreement":null},{"id":"W1914969610","doi":"10.1007/s10664-015-9393-5","title":"An in-depth study of the promises and perils of mining GitHub","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":270,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software; Point (geometry); Data science; World Wide Web; Set (abstract data type); Event (particle physics); Empirical research; The Internet; Quality (philosophy)","score_opus":0.05189227763191294,"score_gpt":0.31378800903066856,"score_spread":0.2618957313987556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1914969610","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85657966,0.015867228,0.06180893,0.027074505,0.00012936456,0.00029159663,0.00198889,0.00022928634,0.036030516],"genre_scores_gemma":[0.94508326,0.00616323,0.042049866,0.0011380754,0.00018276708,0.00007881095,0.0016292557,0.00011313947,0.0035616537],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.99071324,0.005087964,0.00045831472,0.00049070804,0.0028655084,0.00038424777],"domain_scores_gemma":[0.8384849,0.1324434,0.008217656,0.009278706,0.009914197,0.0016611865],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014398425,0.0003185126,0.00037137818,0.004433096,0.0015059564,0.0045652315,0.0013795418,0.000858712,0.0022946477],"category_scores_gemma":[0.10823736,0.0003601127,0.00035523233,0.008001745,0.0023457387,0.010759787,0.0018934367,0.0019939437,0.0005213564],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047844206,0.00064161676,0.27630225,0.0023566908,0.00016052004,0.00032071557,0.016968518,0.0042813,0.0060032434,0.11944281,0.0134437485,0.5596002],"study_design_scores_gemma":[0.00006988324,0.0007331321,0.46164995,0.002946607,0.00020404876,0.0017454573,0.07258553,0.075227104,0.017628208,0.20010546,0.16689067,0.00021411003],"about_ca_topic_score_codex":0.0057446742,"about_ca_topic_score_gemma":0.018905,"teacher_disagreement_score":0.9856016,"about_ca_system_score_codex":0.0018978862,"about_ca_system_score_gemma":0.0038572466,"threshold_uncertainty_score":0.07614708},"labels":[],"label_agreement":null},{"id":"W1916585208","doi":"10.1109/wcre.2000.891468","title":"PBS tool demonstration report on xfig","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Waterloo","funders":"","keywords":"Documentation; Computer science; Java; Source code; Software; Software engineering; Set (abstract data type); Programming language; Turing; Deliverable; Software architecture; Code (set theory); World Wide Web; Engineering; Systems engineering","score_opus":0.03152040929387885,"score_gpt":0.26215635871839227,"score_spread":0.23063594942451343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1916585208","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01813111,0.0007529261,0.32915595,0.0033608566,0.0011265976,0.0012704156,0.044929605,0.40738496,0.19388759],"genre_scores_gemma":[0.088645086,0.0015094632,0.41312212,0.0017115457,0.00036109047,0.0023677244,0.15191856,0.0836175,0.256747],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985568,0.00019134118,0.00009548888,0.00018637549,0.00081744976,0.0001526432],"domain_scores_gemma":[0.9967687,0.0008740029,0.000097880926,0.00070900645,0.0013081407,0.00024221024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027810612,0.0012461453,0.00050395896,0.0015825545,0.0009569538,0.0023420213,0.0022756464,0.0011641058,0.15424134],"category_scores_gemma":[0.008552365,0.00077020476,0.0005946774,0.0017359003,0.00037345936,0.0034699815,0.002293938,0.0016570379,0.07968118],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003345383,0.0001818,0.0021722731,0.00031201716,0.000014935893,0.0004970102,0.00053152273,0.0012547591,0.004644996,0.0049556973,0.68064183,0.3044586],"study_design_scores_gemma":[0.00016437078,0.00015391923,0.0033321895,0.00021970732,0.000018845012,0.00048907776,0.00019504012,0.0058357674,0.011815384,0.004309111,0.973408,0.000058615693],"about_ca_topic_score_codex":0.009993755,"about_ca_topic_score_gemma":0.008890877,"teacher_disagreement_score":0.15424134,"about_ca_system_score_codex":0.00097509066,"about_ca_system_score_gemma":0.0018938358,"threshold_uncertainty_score":0.5159887},"labels":[],"label_agreement":null},{"id":"W1919594131","doi":"10.1109/re.2015.7320422","title":"Inherent characteristics of traceability artifacts less is more","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"National Aeronautics and Space Administration; National Science Foundation","keywords":"Traceability; Computer science; Software engineering","score_opus":0.07959900682346226,"score_gpt":0.30798861490244916,"score_spread":0.2283896080789869,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1919594131","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56990093,0.0037316242,0.39876953,0.0021996237,0.00020497126,0.00022210274,0.0032118629,0.0036630693,0.018096238],"genre_scores_gemma":[0.92840797,0.00083036546,0.0628858,0.000302646,0.000101078695,0.00006386569,0.0034420155,0.00050689856,0.0034594422],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99495476,0.0010067794,0.0005266209,0.0012429978,0.002048093,0.00022071254],"domain_scores_gemma":[0.9537209,0.017071532,0.007938355,0.015380952,0.0052972515,0.00059095194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030091372,0.000785481,0.0006921666,0.004221096,0.00069987593,0.0041766493,0.0011475268,0.0011160972,0.0021434887],"category_scores_gemma":[0.024182223,0.00047534067,0.00096261466,0.0044597425,0.0015333367,0.006936168,0.0016340767,0.001578126,0.0005865339],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037886162,0.0004526892,0.26551744,0.001949516,0.00047518758,0.0015307547,0.0036844742,0.035913385,0.02888311,0.043277573,0.005686478,0.61225057],"study_design_scores_gemma":[0.00007567632,0.0008513452,0.27813175,0.0009098656,0.0007111549,0.0116813015,0.004808576,0.24160008,0.06864411,0.15996082,0.23227884,0.00034653122],"about_ca_topic_score_codex":0.002766196,"about_ca_topic_score_gemma":0.0038205998,"teacher_disagreement_score":0.004221096,"about_ca_system_score_codex":0.00092034735,"about_ca_system_score_gemma":0.00087360945,"threshold_uncertainty_score":0.015914023},"labels":[],"label_agreement":null},{"id":"W1920287393","doi":"10.1002/smr.1695","title":"SCAN: an approach to label and relate execution trace segments","year":2014,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Program comprehension; Computer science; Java; Feature (linguistics); TRACE (psycholinguistics); Precision and recall; Set (abstract data type); Process (computing); Task (project management); Software; Comprehension; Artificial intelligence; Programming language; Data mining; Machine learning; Software system","score_opus":0.015644056639543682,"score_gpt":0.2702961443460459,"score_spread":0.2546520877065022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1920287393","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031142218,0.0003205811,0.9315291,0.00027991572,0.000057064393,0.0006435824,0.0022180725,0.03192217,0.001887226],"genre_scores_gemma":[0.10397516,0.00016148048,0.88632494,0.00013061396,0.000048246857,0.0005513307,0.004182354,0.0018375658,0.0027883598],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9938624,0.0015186642,0.00064691744,0.0015313755,0.0021791756,0.00026149943],"domain_scores_gemma":[0.9750446,0.0111568775,0.0043476317,0.004067681,0.0045402884,0.00084296154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004903202,0.0016444323,0.001289496,0.013022966,0.0013387429,0.0030001057,0.0029042887,0.0021861182,0.006105326],"category_scores_gemma":[0.023508249,0.0008736411,0.0013938885,0.0057978057,0.0013491645,0.0047198567,0.0041722665,0.001623995,0.0025519938],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014180305,0.0005365668,0.02749287,0.0013809794,0.00032767028,0.00068909384,0.008027605,0.011787764,0.054405604,0.016928468,0.022193622,0.8548118],"study_design_scores_gemma":[0.0003386441,0.0013629318,0.021980835,0.0006008644,0.00046487336,0.001777733,0.005584922,0.70311415,0.10699377,0.04783079,0.10955499,0.00039560068],"about_ca_topic_score_codex":0.009148548,"about_ca_topic_score_gemma":0.012902272,"teacher_disagreement_score":0.013022966,"about_ca_system_score_codex":0.0012004732,"about_ca_system_score_gemma":0.0032896665,"threshold_uncertainty_score":0.025930882},"labels":[],"label_agreement":null},{"id":"W1925941278","doi":"10.1002/smr.1635","title":"Detecting asynchrony and dephase change patterns by mining software repositories","year":2013,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal; Lab_Bell (Canada)","funders":"","keywords":"Computer science; Asynchrony (computer programming); Software evolution; Interval (graph theory); Java; Change detection; Software; Data mining; Software development; Programming language; Artificial intelligence","score_opus":0.014671771646186717,"score_gpt":0.258621679834277,"score_spread":0.24394990818809026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1925941278","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9244958,0.0007182926,0.06659531,0.00018939286,0.000025591746,0.00015178673,0.0019256342,0.0050642267,0.0008339129],"genre_scores_gemma":[0.9127042,0.0001910465,0.08125393,0.000042553747,0.000024325069,0.000089994486,0.0051927716,0.00011629592,0.00038484254],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99423814,0.0009151461,0.0009880689,0.0014355041,0.0021821188,0.00024096576],"domain_scores_gemma":[0.96240073,0.01690304,0.009673298,0.005179754,0.005120023,0.0007231443],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004467504,0.00070261286,0.0008838253,0.013725463,0.00056439114,0.001630739,0.0015147799,0.00086995325,0.00030973105],"category_scores_gemma":[0.018250125,0.0004665002,0.0008616852,0.007506015,0.00042044776,0.0020225844,0.001421688,0.00060512044,0.0003116576],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045307213,0.0004542232,0.5378809,0.0005676,0.00044283769,0.00074608345,0.0007180021,0.033670224,0.018180676,0.0010548835,0.003679208,0.40215218],"study_design_scores_gemma":[0.00007312084,0.0003189276,0.20815141,0.0000775923,0.00029692682,0.001239747,0.00056981936,0.7489228,0.032220993,0.003227777,0.0048086005,0.00009218594],"about_ca_topic_score_codex":0.005780811,"about_ca_topic_score_gemma":0.0063333404,"teacher_disagreement_score":0.013725463,"about_ca_system_score_codex":0.00062758446,"about_ca_system_score_gemma":0.0010017735,"threshold_uncertainty_score":0.023626685},"labels":[],"label_agreement":null},{"id":"W1929161383","doi":"10.1002/smr.1659","title":"An empirical study of the effect of file editing patterns on software quality","year":2014,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Computer file; File format; Software; Data file; World Wide Web; Database; Operating system","score_opus":0.01750574317341968,"score_gpt":0.3362430763318027,"score_spread":0.31873733315838304,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1929161383","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99891424,0.000047663678,0.00044904728,0.000070387425,0.00000234695,0.000020770933,0.00012758031,0.000013465446,0.0003544857],"genre_scores_gemma":[0.9994531,0.000019351586,0.000337437,0.000009986807,0.000003717677,0.000014763429,0.00010133812,0.000004230414,0.00005594],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.980682,0.008401194,0.002228272,0.0022359109,0.005214411,0.0012382505],"domain_scores_gemma":[0.426378,0.39068538,0.14453673,0.01439797,0.017067857,0.00693406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012729262,0.00038070732,0.0003297061,0.0031267623,0.00065130176,0.0020371026,0.0013452992,0.0010114537,0.0021214292],"category_scores_gemma":[0.18221885,0.0004092656,0.0006773717,0.0033202542,0.0016238155,0.003502515,0.0016491493,0.0019938422,0.00029797808],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008448503,0.00013985665,0.99447894,0.00002690415,0.00005844979,0.00009658141,0.00071954116,0.00034789424,0.00018024951,0.000057960206,0.00008593963,0.0037232419],"study_design_scores_gemma":[0.000008061892,0.0002634545,0.9945427,0.000015287715,0.000022158463,0.00022429254,0.0017076829,0.00263249,0.00028639557,0.000116499,0.00016519883,0.000015743693],"about_ca_topic_score_codex":0.0033616822,"about_ca_topic_score_gemma":0.0029712063,"teacher_disagreement_score":0.012729262,"about_ca_system_score_codex":0.0008850279,"about_ca_system_score_gemma":0.0007629305,"threshold_uncertainty_score":0.06731957},"labels":[],"label_agreement":null},{"id":"W1929268629","doi":"10.1109/csmr.2000.827322","title":"Towards a quantitative assessment of method replacement","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Implementation; Source code; Programming language; Object-oriented programming; Code (set theory); Set (abstract data type)","score_opus":0.07507042256742712,"score_gpt":0.40381114362427034,"score_spread":0.3287407210568432,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1929268629","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.62225354,0.0018328829,0.3560201,0.00085172406,0.00009982937,0.000292863,0.0015331598,0.0015403343,0.015575568],"genre_scores_gemma":[0.90148246,0.00022181423,0.0964067,0.000083579514,0.000042020405,0.00026496695,0.00069281255,0.00013055938,0.0006750712],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9515905,0.019629257,0.004840206,0.0038564005,0.018973539,0.0011101662],"domain_scores_gemma":[0.5041182,0.3402458,0.08136477,0.033722453,0.03746219,0.0030865108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.038696736,0.0010942536,0.0010081144,0.009861224,0.0007034607,0.0031755387,0.0021877799,0.001939461,0.0016525217],"category_scores_gemma":[0.23614582,0.00065587723,0.0006771959,0.00712388,0.0033231925,0.00769106,0.0035901568,0.0027423955,0.0005911596],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011731428,0.0005700769,0.6247577,0.0013310156,0.00042975685,0.00028101908,0.007271102,0.02416564,0.0335643,0.059057564,0.002565969,0.2448327],"study_design_scores_gemma":[0.0001313889,0.0022099903,0.56701636,0.0005099405,0.0003867351,0.001990068,0.008850058,0.22080731,0.062826194,0.10227678,0.03247643,0.0005187124],"about_ca_topic_score_codex":0.0011116788,"about_ca_topic_score_gemma":0.001046439,"teacher_disagreement_score":0.038696736,"about_ca_system_score_codex":0.0016325776,"about_ca_system_score_gemma":0.0016187994,"threshold_uncertainty_score":0.20465034},"labels":[],"label_agreement":null},{"id":"W1930513839","doi":"10.1109/ccece.2000.849579","title":"Extreme programming: a university team design experience","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Extreme programming; Code refactoring; Deliverable; Computer science; Software engineering; Waterfall model; Test suite; Software development; Documentation; Sequence diagram; Unit testing; Extreme programming practices; Software development process; Use Case Diagram; Test case; Systems engineering; Programming language; Software; Unified Modeling Language; Engineering; Class diagram","score_opus":0.08607000593763067,"score_gpt":0.24702900496389055,"score_spread":0.16095899902625987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1930513839","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46153057,0.0034713217,0.29516312,0.011236678,0.0011071588,0.0005387294,0.00024270541,0.0024755432,0.22423415],"genre_scores_gemma":[0.6468049,0.0034954753,0.18772289,0.002105317,0.0004022291,0.00048055244,0.0005316971,0.0012359151,0.15722097],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99359703,0.0036625,0.00023446337,0.00048985577,0.0014297599,0.000586457],"domain_scores_gemma":[0.9930727,0.001263003,0.0002094278,0.000639537,0.001302373,0.0035130153],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007252303,0.00066660764,0.00041394384,0.00075454015,0.004376993,0.0044143274,0.0020407322,0.0012561202,0.011047216],"category_scores_gemma":[0.0057977256,0.0004616028,0.0006406356,0.0012172764,0.0015100475,0.002176547,0.0045818756,0.0032052503,0.004418395],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069917087,0.005705544,0.010289945,0.0005578236,0.0000768832,0.0033498248,0.07501071,0.011755162,0.011656816,0.04229375,0.075172365,0.763432],"study_design_scores_gemma":[0.00019431663,0.003763809,0.0052878596,0.00040502634,0.000045594305,0.0035531258,0.02779749,0.014123294,0.009614909,0.022929462,0.91211116,0.0001739134],"about_ca_topic_score_codex":0.0007952853,"about_ca_topic_score_gemma":0.0014323327,"teacher_disagreement_score":0.011047216,"about_ca_system_score_codex":0.0019699465,"about_ca_system_score_gemma":0.002888051,"threshold_uncertainty_score":0.038354337},"labels":[],"label_agreement":null},{"id":"W1931878052","doi":"10.1109/issre.2000.885858","title":"Thresholds for object-oriented measures","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cistel Technology (Canada); National Research Council Canada","funders":"","keywords":"Object (grammar); Computer science; Fault (geology); Cognition; Test suite; Object-oriented programming; Artificial intelligence; Reliability engineering; Theoretical computer science; Machine learning; Test case; Psychology; Programming language; Engineering","score_opus":0.042184269044132895,"score_gpt":0.2671482449248177,"score_spread":0.2249639758806848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1931878052","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12777266,0.0029555103,0.7903423,0.0029858577,0.0006117676,0.0009792235,0.0010630211,0.0019886706,0.071301036],"genre_scores_gemma":[0.821763,0.00060214725,0.17137307,0.0008960774,0.00030971624,0.0012705367,0.0004945906,0.0003507576,0.0029401262],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97449255,0.006591309,0.0029663648,0.0041026296,0.010580406,0.0012667774],"domain_scores_gemma":[0.818327,0.13191786,0.015915241,0.019095019,0.011937381,0.002807518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018390622,0.001267851,0.001287872,0.004339719,0.001420165,0.007660084,0.0018865176,0.0025521151,0.0070204465],"category_scores_gemma":[0.17517556,0.0008278326,0.0018473626,0.0036840143,0.005533204,0.012004476,0.0045932946,0.0038516896,0.0013693816],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045730997,0.00024497331,0.029250477,0.00094588357,0.00019560836,0.00036241877,0.0028016225,0.0075158,0.0073856018,0.82021683,0.0042149206,0.12640852],"study_design_scores_gemma":[0.000066102104,0.0004265275,0.016143393,0.00025681194,0.000099521785,0.00072024076,0.00077307,0.02374162,0.0051633343,0.93613243,0.016341968,0.00013505113],"about_ca_topic_score_codex":0.0011414606,"about_ca_topic_score_gemma":0.00039468997,"teacher_disagreement_score":0.018390622,"about_ca_system_score_codex":0.0027370292,"about_ca_system_score_gemma":0.0011843019,"threshold_uncertainty_score":0.09726006},"labels":[],"label_agreement":null},{"id":"W1936045157","doi":"10.1109/wpc.2005.29","title":"REGoLive: Web Site Comprehension with Viewpoints","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada","keywords":"Viewpoints; Computer science; Comprehension; Web site; World Wide Web; Visualization; Graph; Human–computer interaction; The Internet; Programming language; Artificial intelligence; Theoretical computer science","score_opus":0.01238624958627875,"score_gpt":0.2490881400056547,"score_spread":0.23670189041937595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1936045157","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021897493,0.0003247567,0.8515941,0.00075743377,0.00010152315,0.00025206618,0.001222976,0.10378903,0.020060554],"genre_scores_gemma":[0.18868577,0.000703071,0.7782076,0.0004756665,0.00006522494,0.0003453046,0.0038128223,0.017311312,0.010393198],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99887544,0.00041506617,0.000056889992,0.00017493029,0.0003897998,0.00008779806],"domain_scores_gemma":[0.9949582,0.0033926435,0.00017685883,0.0008512604,0.00046788144,0.00015311241],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014229985,0.0014664911,0.00061616296,0.0011118784,0.00041035848,0.0025487912,0.0020990118,0.0016393602,0.015830845],"category_scores_gemma":[0.009660065,0.00077465695,0.0011997346,0.0005275199,0.00088797125,0.004975556,0.0031690188,0.0027242973,0.0046246643],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096898613,0.0006120785,0.0045279562,0.0021368682,0.0001195601,0.0030744679,0.014725649,0.019727143,0.099414885,0.08862282,0.11406375,0.6520059],"study_design_scores_gemma":[0.0005315269,0.0004850168,0.0030352673,0.0008019452,0.00012975984,0.0035158065,0.002943851,0.22776544,0.11289724,0.09545605,0.55210984,0.000328291],"about_ca_topic_score_codex":0.0018072797,"about_ca_topic_score_gemma":0.0026824179,"teacher_disagreement_score":0.015830845,"about_ca_system_score_codex":0.00041907962,"about_ca_system_score_gemma":0.000712807,"threshold_uncertainty_score":0.052959442},"labels":[],"label_agreement":null},{"id":"W1940471252","doi":"10.1109/re.2015.7320442","title":"An enhanced requirements gathering interface for open source software development environments","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Computer science; Interface (matter); Open source; Open source software; Software; Embedding; Software development; Requirements analysis; User interface; Software engineering; Human–computer interaction; Operating system; Artificial intelligence","score_opus":0.07126237347119563,"score_gpt":0.3347037719625611,"score_spread":0.2634413984913655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1940471252","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019011546,0.00007744074,0.9633835,0.00026129695,0.00004604404,0.00055650953,0.00011403609,0.014477138,0.0020725306],"genre_scores_gemma":[0.07015733,0.000075994605,0.9251521,0.00026225552,0.000034130062,0.00041331755,0.00057568075,0.00092148536,0.0024076628],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9914117,0.0037220805,0.0010045109,0.0008488751,0.002723101,0.000289733],"domain_scores_gemma":[0.9778152,0.011541322,0.0012642619,0.005117746,0.0036346123,0.0006268654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008849199,0.0009983834,0.00062950084,0.0016531305,0.00031780527,0.0017525336,0.0021488124,0.0014998213,0.0041994313],"category_scores_gemma":[0.02176275,0.00071856094,0.0009829763,0.00078735896,0.0003965194,0.0039159954,0.002474042,0.001964914,0.002135409],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016697092,0.0017828677,0.0059255613,0.001601874,0.00015035998,0.0015036191,0.00517075,0.0040964917,0.23886111,0.013176214,0.008359947,0.71770144],"study_design_scores_gemma":[0.0009965004,0.003893343,0.02199578,0.0012058587,0.00058350177,0.008563171,0.00148296,0.33047995,0.331804,0.016126595,0.2822754,0.0005929399],"about_ca_topic_score_codex":0.0002508357,"about_ca_topic_score_gemma":0.00022302155,"teacher_disagreement_score":0.008849199,"about_ca_system_score_codex":0.00021572305,"about_ca_system_score_gemma":0.000567657,"threshold_uncertainty_score":0.0467996},"labels":[],"label_agreement":null},{"id":"W194943388","doi":"","title":"Qualities of Relevant Software Documentation: An Industrial Study","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Documentation; Software documentation; Relevance (law); Software; Software peer review; Software development; Technical documentation; Software technical review; Software project management; Software engineering; Computer science; Internal documentation; Personal software process; Resource (disambiguation); World Wide Web; Knowledge management; Engineering management; Software construction; Engineering; Political science","score_opus":0.09871714874419732,"score_gpt":0.3254708372073643,"score_spread":0.22675368846316699,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W194943388","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9974474,0.0002224764,0.00033256007,0.00023634634,0.000004969716,0.000029070141,0.000010546073,0.0000074577956,0.0017091166],"genre_scores_gemma":[0.99875414,0.00021543355,0.0005291847,0.000088966146,0.00001379858,0.000022741226,0.000023158595,0.0000070017413,0.00034558502],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.98377705,0.009913519,0.0011207492,0.0006845901,0.0035165132,0.000987559],"domain_scores_gemma":[0.76484114,0.16915528,0.028463528,0.006109177,0.023464857,0.00796597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020344073,0.00026052826,0.00050078554,0.0049938685,0.0036744287,0.004043806,0.0008935563,0.0012211581,0.0014925288],"category_scores_gemma":[0.13189645,0.0007124559,0.00025515744,0.00385315,0.002918736,0.0031576154,0.00254844,0.0017662655,0.0003111176],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003035542,0.0022468017,0.52825236,0.00045388847,0.00003411238,0.0015479899,0.39797902,0.00018720188,0.0019662709,0.0013420176,0.0015328865,0.064154014],"study_design_scores_gemma":[0.00009581711,0.002919782,0.6218058,0.0003855031,0.000057721412,0.002431288,0.35297355,0.0010334825,0.0014929756,0.0012585648,0.015470058,0.00007550538],"about_ca_topic_score_codex":0.0029213447,"about_ca_topic_score_gemma":0.005785705,"teacher_disagreement_score":0.020344073,"about_ca_system_score_codex":0.0023803092,"about_ca_system_score_gemma":0.0027562971,"threshold_uncertainty_score":0.10759103},"labels":[],"label_agreement":null},{"id":"W1951989532","doi":"10.1109/ccece.2001.933665","title":"Self organizing maps as a tool for software analysis","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Maintainability; Software; Data mining; Software visualization; Self-organizing map; Software metric; Visualization; Software system; Java; Artificial neural network; Artificial intelligence; Software construction; Software engineering; Programming language","score_opus":0.018360846469550433,"score_gpt":0.24959467960412565,"score_spread":0.2312338331345752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1951989532","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018266774,0.0016071091,0.9841537,0.00043948402,0.000164362,0.00013516027,0.00055922504,0.005752839,0.005361423],"genre_scores_gemma":[0.04620317,0.0013231144,0.9474252,0.00014220663,0.00018662568,0.0006869755,0.00091532234,0.00048923196,0.0026282386],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975389,0.001072265,0.00019645522,0.00023564017,0.00089072017,0.00006607163],"domain_scores_gemma":[0.99702245,0.0019438461,0.00018147405,0.0003540059,0.00039931008,0.00009889916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025574383,0.0013564503,0.0011961062,0.0063119982,0.000821865,0.0039789844,0.0017711418,0.001306302,0.00630251],"category_scores_gemma":[0.0069560306,0.00074998406,0.0013272351,0.0052964883,0.0013156459,0.0024376079,0.0019377235,0.0018169406,0.0024786994],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019564267,0.00015946082,0.0015880659,0.0016253126,0.0005728812,0.0007738017,0.0015372442,0.09098772,0.0059510795,0.27226743,0.050796688,0.5735447],"study_design_scores_gemma":[0.000074971445,0.00008605951,0.0019475649,0.000356336,0.00009711671,0.00054887903,0.00046315597,0.4097076,0.004387681,0.42080048,0.1614022,0.0001278628],"about_ca_topic_score_codex":0.002103469,"about_ca_topic_score_gemma":0.0016083332,"teacher_disagreement_score":0.0063119982,"about_ca_system_score_codex":0.00082131743,"about_ca_system_score_gemma":0.0010312024,"threshold_uncertainty_score":0.02108401},"labels":[],"label_agreement":null},{"id":"W1954896785","doi":"10.1109/wpc.1998.693273","title":"Pattern visualization for software comprehension","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software visualization; Program comprehension; Visualization; Reverse engineering; Software engineering; Source code; Documentation; Architectural pattern; Software system; Human–computer interaction; Software design; Programming language; Software development; Software; Software construction; Artificial intelligence","score_opus":0.041449596126794226,"score_gpt":0.2867815646752285,"score_spread":0.24533196854843425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1954896785","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007880675,0.0011998289,0.9759804,0.001841572,0.00008774097,0.000092525996,0.00024301728,0.0069478163,0.0057264497],"genre_scores_gemma":[0.12466385,0.0012706415,0.8705024,0.00019958477,0.00009355742,0.00023370417,0.00036942982,0.0008786451,0.0017881547],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99852103,0.0008298924,0.00009622087,0.00022785463,0.00025638475,0.00006871318],"domain_scores_gemma":[0.9917951,0.0053389026,0.0004187796,0.0015149475,0.00071189133,0.00022040642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022967872,0.0012499719,0.0006977946,0.0031806824,0.0012617512,0.0046291053,0.0010779471,0.0017229018,0.010849487],"category_scores_gemma":[0.017199568,0.00075584016,0.00086524815,0.002110445,0.0017047605,0.006660061,0.0040143006,0.0023907095,0.0014710572],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031106305,0.000077819466,0.002619817,0.0008457941,0.00007517789,0.00031353653,0.005130757,0.01163428,0.019309653,0.242288,0.0280897,0.6893044],"study_design_scores_gemma":[0.00015567245,0.00016533835,0.0026217147,0.0005434702,0.00007892174,0.0012102188,0.0017350543,0.21679996,0.022139318,0.582924,0.17149667,0.00012963133],"about_ca_topic_score_codex":0.0012911666,"about_ca_topic_score_gemma":0.0011030775,"teacher_disagreement_score":0.010849487,"about_ca_system_score_codex":0.00087669527,"about_ca_system_score_gemma":0.00088006194,"threshold_uncertainty_score":0.036295116},"labels":[],"label_agreement":null},{"id":"W1955925169","doi":"10.1109/step.1999.798403","title":"Evidence driven object identification in procedural code","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Porting; Computer science; Legacy system; Programming language; Legacy code; Identification (biology); Software maintenance; Object-oriented programming; Software engineering; Object (grammar); Program comprehension; Software system; Software; Artificial intelligence","score_opus":0.04427552333397929,"score_gpt":0.30792674700575634,"score_spread":0.26365122367177707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1955925169","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047208607,0.00017219713,0.9490723,0.00026585144,0.000022573266,0.00006344132,0.000040717216,0.0015377256,0.0016165976],"genre_scores_gemma":[0.43455106,0.00017658221,0.5612048,0.00012167503,0.00003768476,0.00011182089,0.00018533175,0.0003586399,0.003252492],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99697626,0.0007095397,0.00018670924,0.00040926336,0.0014707698,0.00024753597],"domain_scores_gemma":[0.9755441,0.0155465,0.003074115,0.0032272013,0.0023126407,0.00029543872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035270634,0.00032400002,0.0005813102,0.002531747,0.0008444348,0.002504827,0.0019410305,0.0018445565,0.0018080429],"category_scores_gemma":[0.028651608,0.0007143629,0.00071687053,0.00134582,0.0044190013,0.0033772544,0.0029209657,0.0019198732,0.0007864393],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005655114,0.0001950114,0.014943565,0.00043345412,0.00008555094,0.0010359748,0.0019706236,0.08975402,0.03339296,0.38902614,0.002634906,0.4659623],"study_design_scores_gemma":[0.00003886119,0.00008930631,0.0022090755,0.00008664016,0.00004209212,0.00048988796,0.00024497998,0.6481195,0.038935147,0.30293986,0.0067489133,0.00005573727],"about_ca_topic_score_codex":0.0018509667,"about_ca_topic_score_gemma":0.002249347,"teacher_disagreement_score":0.0035270634,"about_ca_system_score_codex":0.0009843273,"about_ca_system_score_gemma":0.0012189018,"threshold_uncertainty_score":0.018653095},"labels":[],"label_agreement":null},{"id":"W1957789923","doi":"10.1109/icsm.1998.738502","title":"Practices of software maintenance","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":91,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Documentation; Software maintenance; Software engineering; Computer science; Software system; Work (physics); Field (mathematics); Software; Information system; Software development; Scale (ratio); Engineering management; Engineering; Programming language","score_opus":0.04094143050561529,"score_gpt":0.2869531229939031,"score_spread":0.2460116924882878,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1957789923","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8681485,0.005670232,0.040407356,0.010375187,0.00008682791,0.0001475496,0.000054735367,0.00019243582,0.07491723],"genre_scores_gemma":[0.9909384,0.0011101909,0.005190331,0.00023670681,0.000014117333,0.000039007988,0.000027071941,0.000025385501,0.0024187006],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.97404975,0.01786387,0.00078731566,0.0014512611,0.004690379,0.0011574611],"domain_scores_gemma":[0.9436789,0.039723534,0.006927043,0.0044435593,0.0037609206,0.0014660886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009540439,0.00036756156,0.00023206181,0.0032619454,0.004493819,0.0048103966,0.0014279942,0.001651173,0.0016954328],"category_scores_gemma":[0.052811626,0.0005033094,0.00019146762,0.0026518435,0.010744195,0.005331368,0.004391623,0.0018017525,0.00027838055],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000045376357,0.000081192215,0.019995358,0.00023135576,0.00001550055,0.0005521257,0.84065896,0.0004898404,0.002195687,0.041053288,0.0017120646,0.09296931],"study_design_scores_gemma":[0.000025035604,0.00023678047,0.02576573,0.0009300415,0.00003060157,0.002320276,0.6228794,0.0027488647,0.0024698526,0.03905131,0.3034617,0.00008040572],"about_ca_topic_score_codex":0.0056083263,"about_ca_topic_score_gemma":0.0059668897,"teacher_disagreement_score":0.009540439,"about_ca_system_score_codex":0.0048109055,"about_ca_system_score_gemma":0.00394868,"threshold_uncertainty_score":0.050455272},"labels":[],"label_agreement":null},{"id":"W1964411424","doi":"10.1145/2642937.2642981","title":"Dompletion","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs; Intel Corporation","keywords":"Computer science; JavaScript; Scripting language; Document Object Model; Programming language; Code (set theory); Java; Source code; Compiler; Redundant code; Web application; Object (grammar); World Wide Web; Code generation; Operating system; Artificial intelligence; Web page; Set (abstract data type)","score_opus":0.013396580479155996,"score_gpt":0.24704537458730777,"score_spread":0.23364879410815179,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1964411424","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016819732,0.0009265644,0.48045588,0.0004156109,0.00036193922,0.0007381722,0.01036999,0.4718261,0.018086007],"genre_scores_gemma":[0.10043091,0.00062162796,0.78368044,0.0009554307,0.0001592443,0.0007030488,0.04001079,0.048382837,0.025055727],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9969779,0.000587999,0.0002743921,0.0008203727,0.0011750769,0.00016421668],"domain_scores_gemma":[0.99026847,0.004224917,0.0008358581,0.0027574932,0.0016126467,0.00030056734],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.002123841,0.0017381898,0.0008118431,0.0022466574,0.00081690314,0.0023153718,0.0018432671,0.0015210022,0.013117567],"category_scores_gemma":[0.01673287,0.0009090196,0.0012082383,0.0011189124,0.0006564069,0.0033893,0.003000689,0.0022925371,0.016392935],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011196369,0.00031587222,0.009969746,0.0021731611,0.00012875142,0.00096217013,0.001785325,0.0035705378,0.03954226,0.013806624,0.27966243,0.6469635],"study_design_scores_gemma":[0.00013631102,0.0002405366,0.007650469,0.00042914445,0.000083567174,0.0021761986,0.00031866255,0.09699264,0.11031891,0.012261845,0.76914155,0.00025018695],"about_ca_topic_score_codex":0.0010657919,"about_ca_topic_score_gemma":0.0019247304,"teacher_disagreement_score":0.98688245,"about_ca_system_score_codex":0.000502393,"about_ca_system_score_gemma":0.0013710121,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W1964735878","doi":"10.1002/j.2333-8504.2005.tb01990.x","title":"ANALYSIS OF DISCOURSE FEATURES AND VERIFICATION OF SCORING LEVELS FOR INDEPENDENT AND INTEGRATED PROTOTYPE WRITTEN TASKS FOR THE NEW TOEFL®","year":2005,"lang":"en","type":"article","venue":"ETS Research Report Series","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute for Christian Studies; University of Toronto","funders":"","keywords":"Test of English as a Foreign Language; Computer science; Natural language processing; Psychology; Artificial intelligence; Mathematics education; Language assessment","score_opus":0.075963810214805,"score_gpt":0.39649939581488464,"score_spread":0.32053558560007966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1964735878","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9983405,0.00003276816,0.0005350212,0.00001076389,0.0000045056754,0.000043902368,0.00010704889,0.000021539741,0.0009038528],"genre_scores_gemma":[0.9959709,0.000023340894,0.0021090158,0.000011154667,0.000008849248,0.00012138297,0.00030135945,0.00001788547,0.0014359755],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9955914,0.0013545477,0.0007100308,0.00087779795,0.0011951978,0.00027103603],"domain_scores_gemma":[0.89940584,0.06557185,0.010325213,0.006136184,0.016407277,0.0021536315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066179316,0.000700636,0.00049584,0.0036540332,0.000576063,0.0016028472,0.00059357745,0.0006446332,0.0028927922],"category_scores_gemma":[0.061495002,0.000299939,0.00032425777,0.001147506,0.0009802959,0.0010160679,0.001890846,0.0006020396,0.00055114005],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039380896,0.0011697904,0.59408516,0.0005366832,0.0002360457,0.0012977193,0.042212665,0.0009733316,0.1607037,0.00069013296,0.0007542687,0.19340248],"study_design_scores_gemma":[0.000047538455,0.0012401845,0.97717834,0.000025084883,0.000028651744,0.0003675392,0.0035165,0.001064474,0.015397822,0.00020406445,0.00089159404,0.000038271224],"about_ca_topic_score_codex":0.0010869411,"about_ca_topic_score_gemma":0.0017051914,"teacher_disagreement_score":0.0066179316,"about_ca_system_score_codex":0.00065050024,"about_ca_system_score_gemma":0.00037827026,"threshold_uncertainty_score":0.03499937},"labels":[],"label_agreement":null},{"id":"W1964758092","doi":"10.1016/s0304-3975(99)00085-7","title":"Semantic distance between specifications","year":2000,"lang":"en","type":"article","venue":"Theoretical Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; Université Laval","funders":"","keywords":"Computer science; Reuse; Programming language; Semantic similarity; Semantic computing; Software; Software engineering; Theoretical computer science; Information retrieval; Semantic Web; Engineering","score_opus":0.023179762486844023,"score_gpt":0.2687739268575994,"score_spread":0.24559416437075535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1964758092","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12266399,0.00074051763,0.8267023,0.0023125513,0.00027615746,0.00016160482,0.0019043422,0.0014615406,0.043776963],"genre_scores_gemma":[0.7605579,0.0006096805,0.21797357,0.00056342117,0.00016748131,0.00018191806,0.0055714324,0.0006336311,0.01374104],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9920853,0.001787549,0.0009410137,0.0015006647,0.0031995333,0.00048586313],"domain_scores_gemma":[0.98776245,0.005744881,0.0007434978,0.0025174236,0.0027585148,0.00047321653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031126002,0.000720495,0.0008392004,0.0045423973,0.0018001753,0.0051015317,0.0017285501,0.0022034757,0.0088589825],"category_scores_gemma":[0.015249002,0.0011207616,0.0018817658,0.003707973,0.0025692119,0.01407294,0.0046324898,0.0038281223,0.0023768672],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001626038,0.00006775401,0.0012162804,0.00016823562,0.000044013723,0.00019818106,0.0010276636,0.0024293454,0.0028763697,0.94862527,0.0013880976,0.041796297],"study_design_scores_gemma":[0.000027844897,0.000057951453,0.0005879827,0.00006552915,0.00006973048,0.00026321204,0.0005566003,0.011786967,0.005056643,0.9625809,0.018919349,0.000027217306],"about_ca_topic_score_codex":0.001196322,"about_ca_topic_score_gemma":0.0010014881,"teacher_disagreement_score":0.0088589825,"about_ca_system_score_codex":0.0018440306,"about_ca_system_score_gemma":0.001604898,"threshold_uncertainty_score":0.029636264},"labels":[],"label_agreement":null},{"id":"W1964953506","doi":"10.5539/cis.v6n3p68","title":"An Assessment of Changeability of Open Source Software","year":2013,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Maintainability; Computer science; Metric (unit); Software metric; Software; Open source software; Java; Open source; Software engineering; Software system; Coupling (piping); Software development; Data mining; Reliability engineering; Software quality; Programming language","score_opus":0.02490744047426789,"score_gpt":0.33844341992606214,"score_spread":0.31353597945179423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1964953506","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98738897,0.00016467679,0.009465563,0.0000534992,0.000008547358,0.00007185879,0.00015868047,0.00014459752,0.0025435407],"genre_scores_gemma":[0.9945814,0.000053007567,0.0047625285,0.000005657462,0.000006269195,0.000025840354,0.00023115621,0.0000139937965,0.00032015186],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9940783,0.0012103181,0.00050728983,0.00047157056,0.0035304788,0.00020211],"domain_scores_gemma":[0.9506814,0.026724963,0.009558395,0.0024854932,0.009296157,0.0012537084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039356165,0.0005210104,0.00043338788,0.008545834,0.00043020464,0.0013360616,0.00047492058,0.00069744524,0.00081289123],"category_scores_gemma":[0.03716945,0.00021194314,0.0005639591,0.004053056,0.0004686804,0.0021488827,0.0009885295,0.0005009457,0.00021016925],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037660284,0.000501175,0.7945798,0.0003162673,0.00032983447,0.0004784554,0.00274896,0.0117705455,0.031345703,0.0015548846,0.00041656313,0.1555811],"study_design_scores_gemma":[0.000011103802,0.0010350844,0.9366571,0.000050732997,0.00007015477,0.00046951644,0.0011711875,0.049284622,0.009115317,0.0010689936,0.0010056171,0.00006062386],"about_ca_topic_score_codex":0.002146464,"about_ca_topic_score_gemma":0.0020830075,"teacher_disagreement_score":0.008545834,"about_ca_system_score_codex":0.0006484162,"about_ca_system_score_gemma":0.00039658614,"threshold_uncertainty_score":0.020813823},"labels":[],"label_agreement":null},{"id":"W1965199287","doi":"10.1109/ccece.2014.6901018","title":"Near-miss software clones in open source games: An empirical study","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Saskatchewan","funders":"","keywords":"clone (Java method); Cloning (programming); Source code; Computer science; Open source; Programming language; Java; Reuse; Python (programming language); Code reuse; Software; Theoretical computer science; World Wide Web; Engineering; Biology; Genetics","score_opus":0.04120763377302648,"score_gpt":0.35245242958493084,"score_spread":0.31124479581190434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1965199287","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99747616,0.00012406144,0.0014762757,0.000050406932,0.000003292162,0.000055159868,0.00012616014,0.000028511931,0.0006599988],"genre_scores_gemma":[0.9965485,0.00015298478,0.0022182213,0.000051481,0.0000059060962,0.00008732112,0.00038214517,0.000033110566,0.0005203646],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9873187,0.00528366,0.0011199896,0.0018200771,0.0039027187,0.0005548061],"domain_scores_gemma":[0.83994997,0.11215994,0.024208846,0.006767204,0.013859701,0.0030543585],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0080036875,0.00044750012,0.0004514683,0.0038801653,0.0010647639,0.0018689932,0.0012112141,0.0009661295,0.0011097774],"category_scores_gemma":[0.08885765,0.0004245589,0.00039010632,0.0027511234,0.0021383099,0.0034610347,0.0022547278,0.0017589411,0.00043432173],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028943748,0.0013314266,0.8925551,0.00045046513,0.00013575538,0.0011453426,0.044816628,0.0009851946,0.0019352188,0.0016306135,0.0018948786,0.052829962],"study_design_scores_gemma":[0.000040150917,0.0008610308,0.9332397,0.00031012692,0.00010291543,0.0022320463,0.03847782,0.012016529,0.002832926,0.0014913765,0.008296246,0.00009925703],"about_ca_topic_score_codex":0.003511727,"about_ca_topic_score_gemma":0.0060654767,"teacher_disagreement_score":0.9919963,"about_ca_system_score_codex":0.0010807604,"about_ca_system_score_gemma":0.0010570186,"threshold_uncertainty_score":0.042328},"labels":[],"label_agreement":null},{"id":"W1965772655","doi":"10.1109/msr.2013.6624048","title":"The MSR Cookbook: Mining a decade of research","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Best practice; Scope (computer science); Artifact (error); Data science; Categorization; Theme (computing); Computer science; Library science; Political science; Public relations; World Wide Web","score_opus":0.058975194820489484,"score_gpt":0.35196800683504886,"score_spread":0.2929928120145594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1965772655","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16759412,0.29260343,0.20372054,0.110309474,0.013107554,0.00360743,0.11271044,0.006715113,0.08963188],"genre_scores_gemma":[0.26300532,0.1493865,0.44466478,0.016209746,0.0035891435,0.0049970066,0.08513128,0.0045435736,0.02847265],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97816306,0.0070608165,0.0048236507,0.002957304,0.0063606403,0.00063456653],"domain_scores_gemma":[0.83696055,0.095465295,0.017839266,0.01924238,0.027001984,0.0034905125],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.02951352,0.00094147696,0.0011838112,0.047802635,0.0029850195,0.010771924,0.0021363706,0.0016911317,0.00599936],"category_scores_gemma":[0.12537374,0.0011529893,0.0013743169,0.07474379,0.0027812463,0.01587526,0.00495609,0.0030232093,0.0031648024],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021095562,0.00008670947,0.020041984,0.016667526,0.00040219107,0.00066825905,0.04314888,0.00080260634,0.0043873126,0.033141945,0.22377653,0.65666515],"study_design_scores_gemma":[0.00002858964,0.00008611089,0.013310052,0.015398493,0.00020859645,0.0004314355,0.024761464,0.0007687397,0.0017116731,0.017662108,0.9255204,0.000112262554],"about_ca_topic_score_codex":0.0055607543,"about_ca_topic_score_gemma":0.014908353,"teacher_disagreement_score":0.97048646,"about_ca_system_score_codex":0.004317803,"about_ca_system_score_gemma":0.012850524,"threshold_uncertainty_score":0.1560843},"labels":[],"label_agreement":null},{"id":"W1966394892","doi":"10.5555/2486788.2487043","title":"Identifying failure inducing developer pairs within developer networks","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software engineering; Code (set theory); Software; Liberian dollar; Software bug; Software system; Software maintenance; Software development; Source code; Software quality; Software construction; Programming language; Business","score_opus":0.03741111536065112,"score_gpt":0.2666118718198485,"score_spread":0.2292007564591974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1966394892","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89788365,0.00051894074,0.092678174,0.00038588938,0.00003926241,0.00027077066,0.000771764,0.00019405439,0.0072574746],"genre_scores_gemma":[0.9669389,0.00028527333,0.029118795,0.00005172704,0.000024243907,0.00020866405,0.0010717739,0.00004776552,0.002252954],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9956086,0.0015028509,0.00020929528,0.0012004366,0.0009920907,0.00048670595],"domain_scores_gemma":[0.97304046,0.015248413,0.004707306,0.0019814263,0.0034089598,0.001613439],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027479907,0.000619167,0.0005989593,0.006070875,0.001711749,0.0015255277,0.0010403016,0.0012211447,0.0021691828],"category_scores_gemma":[0.03125359,0.0006277846,0.0004177906,0.002796139,0.0007570667,0.0025931653,0.0026334417,0.00079659827,0.0007860561],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042317493,0.0002857368,0.84781224,0.00022249746,0.0001381596,0.0029112666,0.006532646,0.018949509,0.008259199,0.018011956,0.0039051254,0.09254846],"study_design_scores_gemma":[0.00010383363,0.0005013382,0.4021826,0.0002489576,0.0004718234,0.007845103,0.016410194,0.4504454,0.016525924,0.06973157,0.035374824,0.00015842463],"about_ca_topic_score_codex":0.0052008717,"about_ca_topic_score_gemma":0.008929839,"teacher_disagreement_score":0.006070875,"about_ca_system_score_codex":0.001197234,"about_ca_system_score_gemma":0.0011187302,"threshold_uncertainty_score":0.014532983},"labels":[],"label_agreement":null},{"id":"W1966649174","doi":"10.5555/2819009.2819156","title":"Mining temporal properties of data invariants","year":2015,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Correctness; Temporal logic; Computer science; Property (philosophy); Linear temporal logic; Field (mathematics); Theoretical computer science; Data mining; Programming language; Mathematics","score_opus":0.21333505880022077,"score_gpt":0.32497960589375685,"score_spread":0.11164454709353608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1966649174","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35779583,0.00049564085,0.62733567,0.0005214733,0.00004257056,0.0002987458,0.004766681,0.0063471044,0.0023962688],"genre_scores_gemma":[0.7525815,0.000204605,0.23953657,0.000110386856,0.000025749276,0.00025089752,0.005920538,0.00042906628,0.00094073743],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.996847,0.0003753619,0.00034916477,0.0006887718,0.0014908896,0.00024881877],"domain_scores_gemma":[0.97854555,0.010271335,0.0044737286,0.0033755405,0.0030052243,0.00032872517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002187605,0.00064882665,0.0005314595,0.003727158,0.00056881626,0.0011628058,0.0010637298,0.0005498103,0.00077146303],"category_scores_gemma":[0.01978897,0.00060515106,0.0012597007,0.0021353045,0.0011669069,0.0026916733,0.001084261,0.0011389375,0.00027840922],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007237272,0.00064042414,0.2366798,0.0016904618,0.00033294357,0.0024717655,0.0021377746,0.14021152,0.085465625,0.060187288,0.006755907,0.46270272],"study_design_scores_gemma":[0.00005718998,0.00028212543,0.020070806,0.00012334388,0.00014966245,0.0006912242,0.00052468566,0.8225375,0.08916174,0.056757726,0.009566233,0.00007774181],"about_ca_topic_score_codex":0.0049922573,"about_ca_topic_score_gemma":0.008657235,"teacher_disagreement_score":0.0049922573,"about_ca_system_score_codex":0.0012316563,"about_ca_system_score_gemma":0.0027482323,"threshold_uncertainty_score":0.011569321},"labels":[],"label_agreement":null},{"id":"W1966968814","doi":"10.1109/icstw.2013.17","title":"A Call Graph Mining and Matching Based Defect Localization Technique","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Debugging; Software bug; Program slicing; Static analysis; Java; Tree (set theory); Path (computing); Matching (statistics); Source code; Software; Focus (optics); Code (set theory); Call graph; Distributed computing; Theoretical computer science; Programming language; Set (abstract data type)","score_opus":0.009776589453432048,"score_gpt":0.2358316760854896,"score_spread":0.22605508663205756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1966968814","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07504461,0.00050376175,0.9101663,0.0003620251,0.000066875305,0.0004027301,0.00075975974,0.010145711,0.002548257],"genre_scores_gemma":[0.362967,0.00028521195,0.62836945,0.00023058947,0.000036641693,0.00025164778,0.0019548258,0.00041894722,0.005485564],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99840695,0.00012524618,0.00011045455,0.00040869217,0.00080243254,0.00014625846],"domain_scores_gemma":[0.9980317,0.00049101526,0.00039011994,0.0003466164,0.00066457526,0.00007593158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006077993,0.0008997044,0.00092499395,0.005484997,0.00070683524,0.0007631512,0.001884337,0.001401196,0.0018500165],"category_scores_gemma":[0.0030553786,0.00044173084,0.0013947012,0.0030672303,0.0004186372,0.0013508217,0.0009785336,0.0008923907,0.0010608804],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024981756,0.0004975944,0.014013816,0.00036535558,0.00016818856,0.0010041872,0.00033836628,0.030095058,0.07820421,0.005819646,0.009062901,0.8601808],"study_design_scores_gemma":[0.000056296492,0.00034704307,0.011650797,0.000058511403,0.00019535373,0.0029864162,0.00029146177,0.89093846,0.07126876,0.008139332,0.013979264,0.000088302586],"about_ca_topic_score_codex":0.006286398,"about_ca_topic_score_gemma":0.008172613,"teacher_disagreement_score":0.006286398,"about_ca_system_score_codex":0.0005556655,"about_ca_system_score_gemma":0.0017960954,"threshold_uncertainty_score":0.01249963},"labels":[],"label_agreement":null},{"id":"W1967388635","doi":"10.1109/icsme.2014.113","title":"Model Clone Detector Evaluation Using Mutation Analysis","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Computer science; Software engineering; Software maintenance; Precision and recall; Unified Modeling Language; Implementation; Software; Programming language; Data mining; Software system; Machine learning; Gene","score_opus":0.06481390715856372,"score_gpt":0.33678744589302434,"score_spread":0.2719735387344606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1967388635","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5693759,0.0013428563,0.4113547,0.0003623994,0.0001838958,0.0006441833,0.0009829256,0.011218373,0.0045348685],"genre_scores_gemma":[0.7444372,0.00024619675,0.25220013,0.00009666826,0.000023252167,0.00019152332,0.001491059,0.00047320564,0.00084078603],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9796551,0.0062867687,0.0022982536,0.00305292,0.0080009205,0.00070605485],"domain_scores_gemma":[0.9019396,0.058087766,0.0074774395,0.008070604,0.023568252,0.0008563204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019918866,0.001304931,0.0013153696,0.011982362,0.0010606116,0.0031252704,0.00213819,0.00176019,0.00086551375],"category_scores_gemma":[0.10210335,0.00038471838,0.0014258451,0.0037873443,0.0010291198,0.0032415697,0.0020515332,0.0010574597,0.00031057317],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001502775,0.00079254416,0.17039464,0.0011155158,0.0012480235,0.0006809582,0.0019365617,0.1558929,0.059697296,0.014677885,0.0075519937,0.58450896],"study_design_scores_gemma":[0.00012894337,0.00069538935,0.023167105,0.00015129305,0.00035494534,0.0007153987,0.0005685245,0.8720754,0.09321726,0.0045486344,0.004230363,0.00014682884],"about_ca_topic_score_codex":0.008986779,"about_ca_topic_score_gemma":0.007855253,"teacher_disagreement_score":0.019918866,"about_ca_system_score_codex":0.0025768664,"about_ca_system_score_gemma":0.002203804,"threshold_uncertainty_score":0.10534233},"labels":[],"label_agreement":null},{"id":"W1967408872","doi":"10.1109/icpc.2013.6613857","title":"SimCad: An extensible and faster clone detection tool for large scale software systems","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Cloning (programming); Software; Plug-in; Scalability; Software maintenance; Software system; Software engineering; Operating system; Programming language; Biology","score_opus":0.013787678590481134,"score_gpt":0.24775346400999526,"score_spread":0.23396578541951413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1967408872","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010068713,0.0002906162,0.5253673,0.0001523928,0.00008961838,0.0002542786,0.0022648363,0.4596963,0.0018159375],"genre_scores_gemma":[0.09837921,0.00047708047,0.855121,0.0003392587,0.00006686365,0.00076775555,0.010166992,0.028898036,0.005783748],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986714,0.00015295218,0.00014785932,0.00031213165,0.0006201845,0.00009552753],"domain_scores_gemma":[0.99526995,0.0023665319,0.00058815157,0.0008345698,0.00069201563,0.00024872477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022639418,0.001743836,0.0012768722,0.004196728,0.0006152484,0.0015777646,0.0036714936,0.001341825,0.010682725],"category_scores_gemma":[0.009788843,0.0014098515,0.0018698722,0.0023138523,0.0007738475,0.00356676,0.0023707354,0.0017454015,0.0038553418],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012580826,0.00039051994,0.01252844,0.0019253886,0.00046312407,0.00193128,0.001222126,0.03373538,0.07295456,0.012573684,0.15128163,0.70973575],"study_design_scores_gemma":[0.0010334584,0.000572277,0.011393636,0.00036374907,0.00033363514,0.0035250408,0.000247564,0.61689126,0.13920791,0.022613805,0.20325401,0.0005637627],"about_ca_topic_score_codex":0.002537116,"about_ca_topic_score_gemma":0.0028517991,"teacher_disagreement_score":0.010682725,"about_ca_system_score_codex":0.0008081587,"about_ca_system_score_gemma":0.0014253266,"threshold_uncertainty_score":0.035737276},"labels":[],"label_agreement":null},{"id":"W1967828228","doi":"10.1007/s11219-009-9082-y","title":"Improving design-pattern identification: a new approach and an exploratory study","year":2009,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"","keywords":"Identification (biology); Computer science; Constraint (computer-aided design); False positive paradox; Data mining; Resource (disambiguation); Resource constraints; Machine learning; Engineering; Distributed computing","score_opus":0.09743104232199579,"score_gpt":0.355633719165051,"score_spread":0.2582026768430552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1967828228","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.183565,0.00027690033,0.80658805,0.00064250216,0.000053971504,0.0013791277,0.00023396757,0.0013056075,0.005954901],"genre_scores_gemma":[0.2767835,0.0002198894,0.7186904,0.0002578679,0.000027218948,0.000776261,0.0003394391,0.00017068397,0.0027348085],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98791444,0.005494983,0.00090152473,0.001944293,0.003476036,0.0002687038],"domain_scores_gemma":[0.9596115,0.024232216,0.002116673,0.0066988855,0.006864079,0.00047668096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00658013,0.0014080937,0.0013000619,0.0037499212,0.0014604185,0.002963346,0.0024263149,0.0018273193,0.003942651],"category_scores_gemma":[0.038381644,0.0006715159,0.0012138746,0.0024333019,0.0014702387,0.007280332,0.0038932594,0.001860721,0.0010576682],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005452438,0.0039534904,0.037552677,0.0016336852,0.00021513959,0.0005376416,0.013606517,0.003788567,0.045769766,0.0103226,0.0024773742,0.8795973],"study_design_scores_gemma":[0.001729308,0.010616765,0.108052894,0.0014170781,0.002280945,0.010636152,0.049176484,0.4302506,0.22320312,0.0910166,0.071033,0.0005871531],"about_ca_topic_score_codex":0.001407343,"about_ca_topic_score_gemma":0.0030826603,"teacher_disagreement_score":0.00658013,"about_ca_system_score_codex":0.0010933132,"about_ca_system_score_gemma":0.002187964,"threshold_uncertainty_score":0.034799457},"labels":[],"label_agreement":null},{"id":"W1967875050","doi":"10.1145/1159733.1159746","title":"A comparative study of attribute weighting heuristics for effort estimation by analogy","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Analogy; Heuristics; Weighting; Computer science; Estimation; Artificial intelligence; Data mining; Engineering; Epistemology; Philosophy; Systems engineering","score_opus":0.02688793323132183,"score_gpt":0.3138445283460533,"score_spread":0.2869565951147315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1967875050","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34676373,0.0049849832,0.639947,0.00041942715,0.00015174905,0.0004632834,0.00019515929,0.00085189403,0.0062228176],"genre_scores_gemma":[0.70992583,0.00091078837,0.28782248,0.000112703885,0.00005007552,0.00022308645,0.00030841294,0.000077486766,0.00056914776],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9867417,0.008737111,0.000734416,0.00088974036,0.0025096915,0.00038728587],"domain_scores_gemma":[0.8691736,0.11956553,0.0037046764,0.0034298291,0.0036113143,0.0005150546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017113503,0.0010878796,0.001555793,0.005535734,0.00063340156,0.002060546,0.0015097203,0.0010871295,0.0011758187],"category_scores_gemma":[0.08261256,0.0004777268,0.000948393,0.0041931807,0.0007826217,0.0036566071,0.0011654617,0.0012520202,0.00023923116],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014918343,0.0008107471,0.021874944,0.0009805377,0.000688183,0.00009521883,0.0013987755,0.149651,0.0026140292,0.019333394,0.0022447046,0.7988167],"study_design_scores_gemma":[0.00032047878,0.0018612203,0.023253879,0.00033581146,0.0005403525,0.00034391897,0.0010593405,0.92200077,0.0061159953,0.039963406,0.0039732507,0.00023155141],"about_ca_topic_score_codex":0.0031484864,"about_ca_topic_score_gemma":0.0029484243,"teacher_disagreement_score":0.017113503,"about_ca_system_score_codex":0.0015153664,"about_ca_system_score_gemma":0.0013215208,"threshold_uncertainty_score":0.09050596},"labels":[],"label_agreement":null},{"id":"W1967936510","doi":"10.1109/icsm.2011.6080800","title":"Classifying field crash reports for fixing bugs: A case study of Mozilla Firefox","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Mozilla Foundation","keywords":"Crash; Computer science; Software bug; Software; Field (mathematics); Computer security; Operating system","score_opus":0.08592830839498096,"score_gpt":0.3165277803994623,"score_spread":0.23059947200448133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1967936510","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99608576,0.00019906773,0.0025759107,0.00017653371,0.000011608711,0.00014689245,0.00037222542,0.00021314793,0.00021883589],"genre_scores_gemma":[0.9836073,0.00023142638,0.014257529,0.00008689015,0.000025841928,0.00010359261,0.0010606598,0.000096655014,0.000530117],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9952701,0.0015493295,0.00048162957,0.00083396013,0.0014987886,0.0003662282],"domain_scores_gemma":[0.9471751,0.03209997,0.008495289,0.003665138,0.0068131005,0.0017513457],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005143333,0.0010375995,0.0006304345,0.005063017,0.0014087082,0.00071485084,0.0015570608,0.0016660555,0.0005057754],"category_scores_gemma":[0.029073225,0.00054255145,0.000740276,0.002832348,0.0009138092,0.0012848997,0.0008454988,0.0012260333,0.00025290137],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012780909,0.004061152,0.7481338,0.0010982929,0.0006484744,0.01815293,0.024698498,0.012365503,0.038151596,0.0007724184,0.0072195577,0.14341964],"study_design_scores_gemma":[0.00032153845,0.0032218997,0.8605666,0.00022769201,0.00046824006,0.012565147,0.01618563,0.059527647,0.036249354,0.0007439195,0.009626575,0.0002957296],"about_ca_topic_score_codex":0.016211912,"about_ca_topic_score_gemma":0.030285997,"teacher_disagreement_score":0.016211912,"about_ca_system_score_codex":0.0010241625,"about_ca_system_score_gemma":0.0009458799,"threshold_uncertainty_score":0.032235146},"labels":[],"label_agreement":null},{"id":"W1968020518","doi":"10.1007/s11704-013-2394-x","title":"Tag recommendation for open source software","year":2013,"lang":"en","type":"article","venue":"Frontiers of Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Software; Graph; Open source software; Information retrieval; Software engineering; Data mining; World Wide Web; Data science; Theoretical computer science; Programming language","score_opus":0.018763647290405663,"score_gpt":0.27058395518509176,"score_spread":0.2518203078946861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1968020518","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4549611,0.020326035,0.45166764,0.0032026675,0.0026755102,0.00043156525,0.018760165,0.023490787,0.024484582],"genre_scores_gemma":[0.7057997,0.004286729,0.20517965,0.0008840684,0.0009975772,0.0002009867,0.039593045,0.0018308139,0.041227497],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981502,0.0003094008,0.00011584861,0.0002760068,0.0009628973,0.00018562871],"domain_scores_gemma":[0.99246126,0.003208296,0.00037936916,0.001130203,0.0024493905,0.00037143237],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012417578,0.00096508494,0.00091204,0.00732335,0.00091617316,0.001455959,0.0008996688,0.0015928312,0.0038184172],"category_scores_gemma":[0.010004131,0.0003273898,0.0010220245,0.0046381326,0.0002614265,0.0027450232,0.00089705596,0.0009588745,0.003667747],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015781058,0.0004741425,0.02510772,0.00090856553,0.00038763633,0.0005206913,0.00030977017,0.0072233006,0.028034657,0.0048042154,0.13786894,0.79278225],"study_design_scores_gemma":[0.00031397425,0.0010524179,0.055958457,0.0003975173,0.0008642291,0.0020303268,0.00076460245,0.7098293,0.06241944,0.036209527,0.12985289,0.00030732524],"about_ca_topic_score_codex":0.010957205,"about_ca_topic_score_gemma":0.027189393,"teacher_disagreement_score":0.010957205,"about_ca_system_score_codex":0.00074526726,"about_ca_system_score_gemma":0.000884681,"threshold_uncertainty_score":0.021786869},"labels":[],"label_agreement":null},{"id":"W1968155288","doi":"10.1007/s11219-014-9250-6","title":"Empirical analysis of factors affecting confirmation bias levels of software engineers","year":2014,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Confirmation bias; Software; Identification (biology); Metric (unit); Computer science; Software bug; Software engineering; Risk analysis (engineering); Engineering; Psychology; Social psychology; Operations management; Programming language; Business","score_opus":0.11496069501911305,"score_gpt":0.3726311614194243,"score_spread":0.25767046640031127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1968155288","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.998722,0.000058582795,0.00032451237,0.000091068825,0.000003708844,0.000015267246,0.000025301684,0.000008773422,0.00075075135],"genre_scores_gemma":[0.9995227,0.00002107618,0.000258308,0.000018323148,0.000003348886,0.0000063207317,0.000025865666,0.0000027893816,0.000141261],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98340607,0.007709969,0.0021393604,0.0009850311,0.0046378328,0.0011217976],"domain_scores_gemma":[0.43437976,0.42543274,0.084752604,0.011521673,0.033094537,0.010818721],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.018796057,0.00028120118,0.00026249597,0.0029167095,0.00089913496,0.0018840859,0.000820833,0.0011232675,0.0032243165],"category_scores_gemma":[0.22558497,0.0003173257,0.0003516003,0.0014934579,0.0011171615,0.0011754888,0.00093506096,0.0011140597,0.00042321876],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003126813,0.00021995681,0.988655,0.000041963405,0.000043132677,0.00009293914,0.0022557445,0.00010042047,0.00089330936,0.00018879904,0.00010219727,0.0070938403],"study_design_scores_gemma":[0.000030031348,0.0005137597,0.99143124,0.000041275034,0.000070015674,0.000301821,0.0045406786,0.00073609897,0.0015615647,0.0003034323,0.00044444765,0.000025580683],"about_ca_topic_score_codex":0.002632102,"about_ca_topic_score_gemma":0.003292751,"teacher_disagreement_score":0.9812039,"about_ca_system_score_codex":0.0010985704,"about_ca_system_score_gemma":0.0035202256,"threshold_uncertainty_score":0.099404216},"labels":[],"label_agreement":null},{"id":"W1968273866","doi":"10.1016/s1389-1286(00)00011-6","title":"A hybrid model for specifying features and detecting interactions","year":2000,"lang":"en","type":"article","venue":"Computer Networks","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial intelligence; Theoretical computer science","score_opus":0.02483886119907776,"score_gpt":0.2698903878882503,"score_spread":0.24505152668917254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1968273866","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033195363,0.000023817574,0.9938824,0.00007109211,0.000011028411,0.000061881045,0.0002462835,0.0018193654,0.00056454726],"genre_scores_gemma":[0.15404132,0.0001034787,0.839476,0.00023085355,0.0000360853,0.00047646504,0.0011994642,0.0006714842,0.0037648808],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99705195,0.00065462786,0.00032229393,0.0006246398,0.0011261894,0.00022026147],"domain_scores_gemma":[0.9907601,0.0045066117,0.0006609921,0.0025140634,0.0012865314,0.00027167916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026131684,0.0014872215,0.0010433159,0.0020141716,0.0007664874,0.0032113942,0.0038649968,0.0025636214,0.0042232713],"category_scores_gemma":[0.009442938,0.0013743915,0.0025301834,0.0013055961,0.0015117222,0.006470864,0.0026769408,0.0024042444,0.0015929366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094664114,0.00055674074,0.01027683,0.0006492867,0.00040749175,0.0017583804,0.0011938778,0.37723204,0.035291854,0.3801465,0.010269501,0.18127084],"study_design_scores_gemma":[0.000041020132,0.00007157996,0.00024384845,0.000025877882,0.00011957883,0.00021299552,0.00004764108,0.92344004,0.0070788786,0.063250504,0.0054263147,0.000041731757],"about_ca_topic_score_codex":0.0069006192,"about_ca_topic_score_gemma":0.009735088,"teacher_disagreement_score":0.0069006192,"about_ca_system_score_codex":0.0011566894,"about_ca_system_score_gemma":0.0018819011,"threshold_uncertainty_score":0.014128208},"labels":[],"label_agreement":null},{"id":"W1968278615","doi":"10.1109/ase.2011.6100098","title":"Analyzing temporal API usage patterns","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Application programming interface; Reuse; Software; Software engineering; Software development; Programming language; World Wide Web","score_opus":0.045932775836273386,"score_gpt":0.2669862240246316,"score_spread":0.22105344818835823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1968278615","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79019356,0.0011201572,0.19543572,0.00043687448,0.000044596818,0.00016236957,0.004936849,0.0038588413,0.003811125],"genre_scores_gemma":[0.90936613,0.00042485824,0.082354054,0.00006874323,0.00002807349,0.00018940563,0.0059113028,0.00029453146,0.001362982],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99773335,0.0003491186,0.00026590237,0.0005203792,0.00096437684,0.00016683027],"domain_scores_gemma":[0.9902284,0.004032825,0.0019656892,0.0014898715,0.002055115,0.00022813062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014450689,0.00047464474,0.0005348134,0.0049994183,0.00046103395,0.00093528867,0.0008534343,0.0005621786,0.00063098595],"category_scores_gemma":[0.011459551,0.00041890744,0.0007518253,0.0043677227,0.00035617265,0.001813167,0.0007394348,0.00064749207,0.00032882366],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004694382,0.00043060695,0.4196919,0.0007211244,0.00044761947,0.0012744562,0.0021571,0.03171776,0.040365025,0.008544454,0.006243267,0.48793727],"study_design_scores_gemma":[0.000031754047,0.00024207792,0.18264252,0.00012184091,0.00027611753,0.002275705,0.0008922556,0.74465454,0.029373515,0.019609889,0.0197594,0.00012044289],"about_ca_topic_score_codex":0.0075468267,"about_ca_topic_score_gemma":0.011984711,"teacher_disagreement_score":0.0075468267,"about_ca_system_score_codex":0.00051108544,"about_ca_system_score_gemma":0.0008690941,"threshold_uncertainty_score":0.015005767},"labels":[],"label_agreement":null},{"id":"W1968695943","doi":"10.1145/1134744.1134756","title":"Empirical relation between coupling and attackability in software systems:","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software quality; Computer science; Software quality control; Maintainability; Software measurement; Software metric; Software quality analyst; Software; Verification and validation; Software sizing; Software system; Software construction; Software development; Software engineering; Reliability engineering; Engineering; Operating system","score_opus":0.03584240837713559,"score_gpt":0.3063305256642189,"score_spread":0.27048811728708333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1968695943","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9916958,0.00043743395,0.0061192755,0.00014832284,0.0000073435535,0.000022196198,0.000082078295,0.000038672628,0.0014489256],"genre_scores_gemma":[0.9985098,0.000098791905,0.0011201664,0.000015347758,0.000009111302,0.000013766749,0.00009345209,0.000011921017,0.00012745959],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9887888,0.004857992,0.0012718115,0.0011188651,0.0035553644,0.0004070878],"domain_scores_gemma":[0.6464457,0.26279533,0.06610042,0.011229288,0.009909524,0.0035196482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007036087,0.00059725135,0.00036253306,0.0051268144,0.00044699493,0.0013167277,0.0007016951,0.0009233332,0.002016501],"category_scores_gemma":[0.1259105,0.00041335946,0.00085564464,0.0038664753,0.0018419999,0.0027234447,0.0020749725,0.0020575507,0.00044801357],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000068324545,0.00008130347,0.98874223,0.00005599893,0.00022300502,0.000060182447,0.00054606254,0.0011443621,0.0005898451,0.00035327478,0.00005779609,0.008077562],"study_design_scores_gemma":[0.000007947404,0.0002613585,0.989911,0.000033898727,0.00008391199,0.00048573056,0.00046607768,0.0065932,0.00065641984,0.001153581,0.00032935973,0.000017432883],"about_ca_topic_score_codex":0.0019375834,"about_ca_topic_score_gemma":0.0012338156,"teacher_disagreement_score":0.007036087,"about_ca_system_score_codex":0.0006521173,"about_ca_system_score_gemma":0.00047447352,"threshold_uncertainty_score":0.037210822},"labels":[],"label_agreement":null},{"id":"W1969265968","doi":"10.1007/s10664-007-9054-4","title":"Analysis of attribute weighting heuristics for analogy-based software effort estimation method AQUA+","year":2007,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Weighting; Heuristics; Data mining; Set (abstract data type); Computer science; Selection (genetic algorithm); Estimation; Exploit; A-weighting; Machine learning; Engineering","score_opus":0.02959484286795332,"score_gpt":0.3421866429248011,"score_spread":0.3125918000568478,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969265968","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4599824,0.00067774014,0.5317091,0.0002533361,0.0000602984,0.0003554786,0.0003126321,0.0017476715,0.0049013845],"genre_scores_gemma":[0.7402712,0.00008547048,0.25808296,0.000076376105,0.000013528166,0.00019341306,0.00038241557,0.000116999116,0.00077764667],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99480057,0.0033500695,0.00028362448,0.00050095486,0.0008875476,0.00017727335],"domain_scores_gemma":[0.95987296,0.033911604,0.0009830343,0.0021178694,0.002872254,0.0002422946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007072041,0.00047825102,0.000793784,0.001614667,0.00061681104,0.0013645452,0.0014060664,0.0008045983,0.0030922838],"category_scores_gemma":[0.045704264,0.0003029659,0.0004887783,0.0016976771,0.00030541996,0.0021327285,0.0010838402,0.0009902479,0.000348295],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013759632,0.0009320373,0.029032068,0.00046869626,0.0002281688,0.00007973331,0.00095299736,0.05086263,0.0073087555,0.014916794,0.0038357826,0.89000636],"study_design_scores_gemma":[0.00012724378,0.00048681474,0.0104635,0.00005040032,0.0001631826,0.00013901581,0.00036223777,0.96813345,0.005691226,0.012261419,0.002078482,0.000043122218],"about_ca_topic_score_codex":0.002589422,"about_ca_topic_score_gemma":0.0030387957,"teacher_disagreement_score":0.007072041,"about_ca_system_score_codex":0.0007429414,"about_ca_system_score_gemma":0.0013832656,"threshold_uncertainty_score":0.03740102},"labels":[],"label_agreement":null},{"id":"W1969366857","doi":"10.1109/icpc.2010.46","title":"Studying the Impact of Social Structures on Software Quality","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software quality; Quality (philosophy); Software; Data science; Source code; Process (computing); Software metric; Software development; Product (mathematics); Code review; Software bug; Software engineering; Data mining","score_opus":0.04254294862608747,"score_gpt":0.3741402883358084,"score_spread":0.33159733970972094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969366857","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9890984,0.0001611982,0.0077984715,0.0006158538,0.000007080219,0.000024898518,0.00011007303,0.000032460415,0.0021514485],"genre_scores_gemma":[0.99841094,0.00010862055,0.0011702687,0.000022018936,0.000016229444,0.000015506803,0.000056825098,0.000008761001,0.00019079998],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99538225,0.002596862,0.00014136496,0.0004348135,0.0010884949,0.000356216],"domain_scores_gemma":[0.8362156,0.12428923,0.02829916,0.0046225106,0.004020406,0.0025531165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051394184,0.00062156375,0.00046578012,0.0030653558,0.0008817931,0.0020719839,0.00057232176,0.0008546114,0.0017486222],"category_scores_gemma":[0.05908087,0.00039357127,0.00063717755,0.0021930127,0.0016764611,0.0035325144,0.0016464666,0.0011164816,0.00028315652],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002627944,0.0007132951,0.8828423,0.00017438448,0.00063302356,0.00030901164,0.002850068,0.030426256,0.0029263427,0.010256337,0.0007104889,0.06789565],"study_design_scores_gemma":[0.000031237694,0.00083856744,0.81077296,0.00006223345,0.00029770494,0.00020253562,0.002459065,0.15925975,0.0024264406,0.022103533,0.0014636302,0.0000823576],"about_ca_topic_score_codex":0.0053377897,"about_ca_topic_score_gemma":0.0075483588,"teacher_disagreement_score":0.0053377897,"about_ca_system_score_codex":0.0018127321,"about_ca_system_score_gemma":0.0009565677,"threshold_uncertainty_score":0.027180195},"labels":[],"label_agreement":null},{"id":"W1969427878","doi":"10.1109/ms.2003.1207449","title":"Inspection's role in software quality assurance","year":2003,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Software quality analyst; Quality assurance; Software quality; Software quality control; Software engineering; Software quality assurance; Computer science; Software security assurance; Software; Software inspection; Verification and validation; Quality (philosophy); Software peer review; Software development; Software construction; Engineering; Reliability engineering; Operations management; Programming language; Computer security; Information security","score_opus":0.020072206611204773,"score_gpt":0.2847007827377494,"score_spread":0.2646285761265446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969427878","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.274772,0.021221086,0.38423002,0.053436406,0.0027819711,0.000107595544,0.00010082856,0.002425264,0.26092488],"genre_scores_gemma":[0.9655408,0.0015278179,0.020500375,0.0010072963,0.00030343444,0.000024344243,0.00001816874,0.00009805662,0.01097981],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9924843,0.0033668082,0.0002636819,0.0006590262,0.0027744612,0.00045174218],"domain_scores_gemma":[0.9311208,0.047945082,0.0052466425,0.0046107005,0.008570911,0.0025058265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012183322,0.00049240945,0.0004872123,0.0020977776,0.001974538,0.0043622125,0.001511132,0.002670682,0.004402108],"category_scores_gemma":[0.057870258,0.0005003623,0.00041012224,0.0014707961,0.0062029557,0.005848757,0.001569265,0.0024079643,0.000652287],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057440466,0.00035032927,0.019755434,0.0007008396,0.00005174522,0.00037325377,0.004218763,0.0064290003,0.008197429,0.42569602,0.014804564,0.5188483],"study_design_scores_gemma":[0.00026054468,0.0016480581,0.058896493,0.0013616229,0.0002665692,0.0029651597,0.0044366047,0.07670186,0.015787859,0.6752417,0.16218454,0.00024899168],"about_ca_topic_score_codex":0.003456157,"about_ca_topic_score_gemma":0.0028576045,"teacher_disagreement_score":0.012183322,"about_ca_system_score_codex":0.0019572377,"about_ca_system_score_gemma":0.00432099,"threshold_uncertainty_score":0.06443232},"labels":[],"label_agreement":null},{"id":"W1969441449","doi":"10.1115/detc2010-29057","title":"Supporting Biomimetic Design by Embedding Metadata in Natural-Language Corpora","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Metadata; Computer science; Natural language; Process (computing); Identification (biology); Information retrieval; Natural language processing; Word embedding; Biological database; Natural (archaeology); Word (group theory); Lexicographical order; Artificial intelligence; Embedding; World Wide Web; Linguistics; Programming language","score_opus":0.016812086828916918,"score_gpt":0.315404339313351,"score_spread":0.29859225248443405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969441449","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06992591,0.0026813226,0.8191799,0.0046237246,0.00042585735,0.001589116,0.04698306,0.034504812,0.020086383],"genre_scores_gemma":[0.09011995,0.0012519689,0.8456399,0.0003279722,0.00012670994,0.0009946545,0.057456225,0.0019910675,0.0020916127],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99180216,0.0038352613,0.0016410783,0.0012453314,0.0013393017,0.00013686567],"domain_scores_gemma":[0.9080577,0.07082158,0.0049521835,0.009220382,0.006050479,0.00089773984],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01028353,0.0015518425,0.0015876656,0.014824356,0.0018822171,0.004928569,0.0022401966,0.0017081553,0.008775197],"category_scores_gemma":[0.061687734,0.0013216849,0.0014105969,0.009858175,0.0018048845,0.016762715,0.005377029,0.0020225036,0.0049406663],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010471635,0.0012744167,0.012753104,0.016885638,0.00040699256,0.0025488762,0.013981759,0.021095406,0.06978806,0.08390615,0.06357727,0.71273506],"study_design_scores_gemma":[0.0005390719,0.00056172453,0.00889731,0.0029340424,0.00051608024,0.0023128064,0.0096990485,0.21125972,0.07670262,0.16717884,0.5187002,0.0006984817],"about_ca_topic_score_codex":0.0036983765,"about_ca_topic_score_gemma":0.008751123,"teacher_disagreement_score":0.014824356,"about_ca_system_score_codex":0.0016789264,"about_ca_system_score_gemma":0.004025947,"threshold_uncertainty_score":0.054385126},"labels":[],"label_agreement":null},{"id":"W1969971118","doi":"10.1016/j.jss.2014.11.040","title":"Investigating the effect of “defect co-fix” on quality assurance resource allocation: A search-based approach","year":2014,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Manitoba","funders":"","keywords":"Artifact (error); Computer science; Quality assurance; Software; Software quality assurance; Resource allocation; Software bug; Quality (philosophy); Code (set theory); Software quality; Source code; Resource (disambiguation); Software inspection; Data mining; Software engineering; Software development; Artificial intelligence; Engineering; Programming language; Operations management","score_opus":0.027913264131770945,"score_gpt":0.2912585648168751,"score_spread":0.2633453006851042,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969971118","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9778535,0.00039678818,0.018421447,0.00039223034,0.00003347551,0.00014415174,0.000104076054,0.00020371158,0.0024505798],"genre_scores_gemma":[0.9924717,0.000045505392,0.0069509842,0.000051254025,0.0000067124047,0.000043870627,0.00004096777,0.0000149032585,0.00037399586],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9942887,0.003591235,0.00024362521,0.00058170356,0.0008136709,0.00048109752],"domain_scores_gemma":[0.86099774,0.12577863,0.005513665,0.0033130688,0.0033281345,0.0010686795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0099191,0.0007891693,0.0014113191,0.0015510161,0.000570481,0.0013319426,0.002311004,0.0020162554,0.003421878],"category_scores_gemma":[0.05380825,0.00047521477,0.0009957564,0.0014204673,0.0011081257,0.0018577857,0.0008708817,0.0013006013,0.00018959281],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.023021482,0.006488247,0.049125012,0.0010788961,0.0019929337,0.00025926763,0.00053736387,0.7332192,0.017748866,0.008876314,0.0014978516,0.15615454],"study_design_scores_gemma":[0.00105883,0.008797146,0.028147765,0.00005673546,0.0014981133,0.0001099813,0.00063784467,0.94681525,0.007946808,0.004282348,0.00057805574,0.000071163566],"about_ca_topic_score_codex":0.011469321,"about_ca_topic_score_gemma":0.009417913,"teacher_disagreement_score":0.011469321,"about_ca_system_score_codex":0.0017083365,"about_ca_system_score_gemma":0.0035313063,"threshold_uncertainty_score":0.05245787},"labels":[],"label_agreement":null},{"id":"W1969987057","doi":"10.1145/986537.986569","title":"Generalization for component reuse","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Regina","funders":"","keywords":"Computer science; Generalization; Correctness; Reuse; Programming language; Abstraction; Component (thermodynamics); Semantics (computer science); Component-based software engineering; Theoretical computer science; Reusability; Formal specification; Software; Software development; Mathematics","score_opus":0.026377697835569158,"score_gpt":0.2792409485633301,"score_spread":0.25286325072776095,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969987057","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008568996,0.00074040995,0.9774993,0.0005278646,0.000054486085,0.00011225085,0.000041444622,0.0011008387,0.01135443],"genre_scores_gemma":[0.25606585,0.001640066,0.72991484,0.0006959142,0.00016370227,0.00038603944,0.00028788185,0.00073053903,0.010115093],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9963176,0.0008549824,0.00034913732,0.001129115,0.0011286649,0.00022036956],"domain_scores_gemma":[0.994125,0.0016863407,0.00036509134,0.0030809243,0.00060756045,0.0001350153],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032351944,0.00091840397,0.00094810984,0.0015245082,0.0014841387,0.0021845584,0.0013038297,0.0016279931,0.0053170975],"category_scores_gemma":[0.008389961,0.00057910255,0.0023987282,0.001345695,0.0047339355,0.007377853,0.004151624,0.0026819676,0.001705466],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003378072,0.000016997585,0.0005138647,0.00017406719,0.000029123912,0.00011187038,0.0003677125,0.009465521,0.0040250914,0.90599906,0.0018277568,0.07743519],"study_design_scores_gemma":[0.000027081976,0.0000649334,0.0004398844,0.00013493768,0.00005953784,0.0005930061,0.000112809845,0.038468093,0.008775236,0.85580444,0.095474586,0.000045450724],"about_ca_topic_score_codex":0.0021023112,"about_ca_topic_score_gemma":0.0014355362,"teacher_disagreement_score":0.0053170975,"about_ca_system_score_codex":0.0018546405,"about_ca_system_score_gemma":0.0014616478,"threshold_uncertainty_score":0.017787457},"labels":[],"label_agreement":null},{"id":"W1970095746","doi":"10.1145/1985441.1985469","title":"Supporting software history exploration","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Source code; Interface (matter); World Wide Web; User interface; Software; Software engineering; Software development; Code (set theory); Human–computer interaction; Data science; Programming language; Operating system","score_opus":0.08173768313967397,"score_gpt":0.27668652507768327,"score_spread":0.1949488419380093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1970095746","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11181668,0.00073001214,0.818426,0.001184089,0.000060704177,0.00034138918,0.0005997221,0.057172496,0.0096688755],"genre_scores_gemma":[0.6015103,0.00050566037,0.38809845,0.0002744833,0.000049958908,0.0003002931,0.0013010011,0.0030239963,0.0049357894],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9953701,0.0021247389,0.0003046224,0.0007955941,0.0011537819,0.00025117607],"domain_scores_gemma":[0.9311474,0.05264911,0.002545638,0.009776765,0.0027080788,0.0011730277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059634903,0.0013783519,0.0010531368,0.0024081543,0.00072153215,0.0039144666,0.0022116967,0.001760924,0.005498156],"category_scores_gemma":[0.06360051,0.0011882812,0.00095516886,0.0014078148,0.0010587277,0.012325252,0.0058345953,0.0014448899,0.0015398404],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018925472,0.0008001628,0.026266636,0.0018606873,0.0001591581,0.001108933,0.0184874,0.017955227,0.03905398,0.054075487,0.02372848,0.8146114],"study_design_scores_gemma":[0.00043976374,0.00093401404,0.008010728,0.00074766995,0.00026240316,0.001776408,0.0041901944,0.6521444,0.05130804,0.13966829,0.14010914,0.0004089572],"about_ca_topic_score_codex":0.0014668677,"about_ca_topic_score_gemma":0.0030739885,"teacher_disagreement_score":0.0059634903,"about_ca_system_score_codex":0.0006355743,"about_ca_system_score_gemma":0.0014880182,"threshold_uncertainty_score":0.031538308},"labels":[],"label_agreement":null},{"id":"W1970652823","doi":"10.1145/1808901.1808908","title":"Finding similar defects using synonymous identifier retrieval","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Japan Society for the Promotion of Science; University of Waterloo","keywords":"Identifier; Fragment (logic); Computer science; Code (set theory); Source code; clone (Java method); Data mining; Information retrieval; Programming language; Biology; Set (abstract data type); Gene; Genetics","score_opus":0.029370276819140072,"score_gpt":0.29739193016820387,"score_spread":0.2680216533490638,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1970652823","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27094096,0.0023238359,0.7059316,0.0006098515,0.00021319337,0.0005211523,0.0015138788,0.01180588,0.006139665],"genre_scores_gemma":[0.44225696,0.00093828625,0.546343,0.00037413574,0.00017426143,0.0002717622,0.0051013837,0.00070447987,0.0038357906],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99733436,0.00039948724,0.0003212932,0.00063136633,0.0011361566,0.00017743511],"domain_scores_gemma":[0.9943356,0.0018930255,0.00090714556,0.0011169154,0.0015458912,0.00020141063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001544113,0.0007911583,0.0017227869,0.012211347,0.0011337516,0.0017360424,0.0020873602,0.0019294902,0.0021230977],"category_scores_gemma":[0.010684786,0.0003879858,0.0011694785,0.006343527,0.0008399448,0.0041281222,0.002756065,0.00065476244,0.0012799187],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006109508,0.0003356612,0.037056725,0.0010414501,0.00031405984,0.0029126743,0.0021255347,0.004797725,0.17028606,0.0112580275,0.01114429,0.7581168],"study_design_scores_gemma":[0.00048005334,0.0018988309,0.08137543,0.00034787573,0.0015882849,0.0277526,0.004086996,0.43327937,0.31957582,0.043565825,0.0852891,0.00075987534],"about_ca_topic_score_codex":0.0031131033,"about_ca_topic_score_gemma":0.0034136483,"teacher_disagreement_score":0.012211347,"about_ca_system_score_codex":0.0007410937,"about_ca_system_score_gemma":0.0014207946,"threshold_uncertainty_score":0.008166194},"labels":[],"label_agreement":null},{"id":"W1971089164","doi":"10.1109/ssbse.2010.26","title":"Concept Location with Genetic Algorithms: A Comparison of Four Distributed Architectures","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Scalability; Computation; Context (archaeology); Genetic algorithm; TRACE (psycholinguistics); Parallel computing; Theoretical computer science; Algorithm; Distributed computing; Machine learning","score_opus":0.015088271099442462,"score_gpt":0.2737144994804967,"score_spread":0.25862622838105426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1971089164","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7521177,0.0022069428,0.22656609,0.00062351517,0.00014880286,0.00037490716,0.00015162188,0.003393016,0.014417309],"genre_scores_gemma":[0.8018394,0.0006869076,0.1947994,0.000070344824,0.000030416917,0.00025339273,0.00023524162,0.00016328218,0.0019216414],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99893564,0.00031637595,0.000058814843,0.00018194942,0.00040607987,0.0001011506],"domain_scores_gemma":[0.99682754,0.0016745675,0.00013790841,0.00055699615,0.0006364324,0.00016649754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018701081,0.0007751878,0.0008594504,0.0017664311,0.000680783,0.0014529742,0.0016289807,0.001468196,0.0015576124],"category_scores_gemma":[0.0050862515,0.000299951,0.0005753491,0.002173519,0.0008660409,0.0014951185,0.0009402999,0.0010404661,0.00030947605],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001792195,0.0009126197,0.0059009055,0.00030546813,0.00025996813,0.00014825836,0.00023766534,0.62383914,0.009765419,0.010519009,0.0015246368,0.3447948],"study_design_scores_gemma":[0.00041613233,0.00064244115,0.002634526,0.000028019385,0.00010856548,0.0000724609,0.00017507975,0.977999,0.008520335,0.0068629244,0.0025098315,0.000030693842],"about_ca_topic_score_codex":0.0075469376,"about_ca_topic_score_gemma":0.005379228,"teacher_disagreement_score":0.0075469376,"about_ca_system_score_codex":0.0019349811,"about_ca_system_score_gemma":0.0016059116,"threshold_uncertainty_score":0.015006006},"labels":[],"label_agreement":null},{"id":"W1971146998","doi":"10.1007/s11219-015-9271-9","title":"On the use of design defect examples to detect model refactoring opportunities","year":2015,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Code refactoring; Computer science; Software bug; Set (abstract data type); Context (archaeology); Similarity (geometry); Data mining; Software; Machine learning; Reliability engineering; Artificial intelligence; Engineering; Programming language","score_opus":0.5757960133181265,"score_gpt":0.376091517068576,"score_spread":0.19970449624955044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1971146998","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.706314,0.0010182303,0.27146843,0.0012117048,0.000102684884,0.0003657509,0.0008565382,0.005059336,0.013603247],"genre_scores_gemma":[0.8073189,0.00028933867,0.18964007,0.00017610003,0.000020242784,0.00004294398,0.00061800535,0.00025896088,0.0016355434],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99556476,0.0013715901,0.00029585842,0.000572738,0.0019613125,0.00023364653],"domain_scores_gemma":[0.8982179,0.08204725,0.0039110514,0.0070357914,0.007878882,0.0009091658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033243657,0.0006584266,0.00058569515,0.004504628,0.0006319375,0.0017216933,0.0016189297,0.0021355879,0.0016625315],"category_scores_gemma":[0.0509396,0.00047051595,0.0005198001,0.0013559117,0.00068413746,0.0023741752,0.0014346903,0.0010448517,0.00053822284],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021972863,0.0014121564,0.16460659,0.0006672396,0.00032796408,0.002227332,0.001885288,0.042310324,0.05122132,0.009788552,0.0069095665,0.71644646],"study_design_scores_gemma":[0.00015290211,0.00060714415,0.043815624,0.00027602026,0.00032077043,0.0025585026,0.0007563421,0.8865381,0.04831105,0.010284001,0.0062337257,0.00014589312],"about_ca_topic_score_codex":0.0059739584,"about_ca_topic_score_gemma":0.01222961,"teacher_disagreement_score":0.0059739584,"about_ca_system_score_codex":0.0004486651,"about_ca_system_score_gemma":0.00081034366,"threshold_uncertainty_score":0.017581105},"labels":[],"label_agreement":null},{"id":"W1971429192","doi":"10.1007/s10009-002-0090-5","title":"A framework for distributing object-oriented designs","year":2003,"lang":"en","type":"article","venue":"International Journal on Software Tools for Technology Transfer","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"Computer science; Object-oriented programming; Object-oriented design; Object (grammar); Software engineering; Theory of computation; Process (computing); Software; Programming language; Distributed computing; Artificial intelligence","score_opus":0.041865767254532116,"score_gpt":0.32452154465712285,"score_spread":0.28265577740259074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1971429192","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008008505,0.000059677015,0.99258125,0.00009490253,0.00004104926,0.00010695576,0.000039232058,0.0047065206,0.0015694739],"genre_scores_gemma":[0.019215459,0.0002357747,0.9730997,0.00007169782,0.000041445364,0.00028946844,0.0002505841,0.0012292517,0.005566637],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99662256,0.00078543974,0.00041122283,0.00040406495,0.0014984536,0.00027822246],"domain_scores_gemma":[0.9955983,0.0010325515,0.00026539864,0.002160548,0.0006472098,0.00029599143],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061098617,0.0016644694,0.0013286929,0.003050943,0.0025384005,0.0065946337,0.004250171,0.002873995,0.010840183],"category_scores_gemma":[0.009523557,0.0022693942,0.0028334497,0.002547892,0.003195753,0.0079909125,0.0049585314,0.003571841,0.0050651478],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010551878,0.00017726484,0.0007911012,0.00030119938,0.00007664994,0.0003130606,0.0009973636,0.01760406,0.005567786,0.7470789,0.010228608,0.21675858],"study_design_scores_gemma":[0.0001654674,0.00014291714,0.00027609096,0.0003294004,0.00020653514,0.0005222798,0.00023435659,0.15117064,0.009676495,0.5920823,0.24506518,0.0001283482],"about_ca_topic_score_codex":0.004282657,"about_ca_topic_score_gemma":0.0060493336,"teacher_disagreement_score":0.010840183,"about_ca_system_score_codex":0.0015802702,"about_ca_system_score_gemma":0.0030608794,"threshold_uncertainty_score":0.036264002},"labels":[],"label_agreement":null},{"id":"W1971867375","doi":"10.1145/1449764.1449786","title":"The impact of static-dynamic coupling on remodularization","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Static analysis; Modular design; Java; Coupling (piping); Programming language; Process (computing); Point (geometry); Mathematics; Engineering","score_opus":0.017808887984384306,"score_gpt":0.2960869992134959,"score_spread":0.2782781112291116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1971867375","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95816934,0.0002209999,0.036918163,0.00018943402,0.000011149263,0.00008211959,0.000031564956,0.00023867677,0.004138424],"genre_scores_gemma":[0.99247277,0.00004869021,0.0071060103,0.000028438613,0.0000055954583,0.000022814396,0.00002748228,0.000067306646,0.00022093386],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9860865,0.0058651394,0.00087868545,0.0013422025,0.004533508,0.0012940881],"domain_scores_gemma":[0.7646028,0.18155727,0.021539075,0.018824616,0.010469454,0.003006812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013436641,0.00069262454,0.00036014154,0.001824052,0.00092603185,0.0021188315,0.001033027,0.0006845908,0.0021395811],"category_scores_gemma":[0.109189555,0.00048391504,0.00052934856,0.00095288147,0.0023620517,0.0041605504,0.002685494,0.0012405337,0.00018354256],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014049091,0.0011555353,0.4141397,0.0012774544,0.00046424652,0.0021318477,0.0107854055,0.12057425,0.11187805,0.035719924,0.000706624,0.29976204],"study_design_scores_gemma":[0.00018402272,0.005926206,0.4835374,0.00045700304,0.00092427846,0.0040047755,0.010283261,0.2700965,0.16252454,0.050790712,0.010864074,0.0004072575],"about_ca_topic_score_codex":0.0015531413,"about_ca_topic_score_gemma":0.0019322628,"teacher_disagreement_score":0.013436641,"about_ca_system_score_codex":0.0010331853,"about_ca_system_score_gemma":0.0011892414,"threshold_uncertainty_score":0.0710606},"labels":[],"label_agreement":null},{"id":"W1972028939","doi":"10.1145/568760.568837","title":"Recovering software requirements from system-user interaction traces","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Documentation; Business process reengineering; User requirements document; Software engineering; Task (project management); User interface; Software requirements specification; Software system; Process (computing); Software; Software development; Software design; Systems engineering; Programming language; Engineering","score_opus":0.0425035807625342,"score_gpt":0.26825602805568466,"score_spread":0.22575244729315047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972028939","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5570261,0.00024187869,0.43184185,0.00035857805,0.000026112188,0.00039442535,0.0030170446,0.005258012,0.0018359898],"genre_scores_gemma":[0.75068444,0.0002189245,0.24077366,0.000044648645,0.000014731061,0.000312008,0.0066455803,0.00030918117,0.0009967838],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963503,0.0010540122,0.0003536885,0.00039358318,0.0016572771,0.00019114847],"domain_scores_gemma":[0.97221464,0.014438088,0.0022694608,0.005342256,0.0052587576,0.0004767919],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025801887,0.0012571417,0.00094077224,0.0045175143,0.0004941287,0.0012715833,0.001445983,0.0013022728,0.0007576615],"category_scores_gemma":[0.03753817,0.0007789369,0.00079907896,0.0021778469,0.00041907886,0.0019096978,0.0014384441,0.0014342525,0.0008673738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008488432,0.00095818075,0.1236577,0.00093398895,0.00023287667,0.0021891359,0.002952104,0.12799443,0.03899083,0.007241296,0.0048425654,0.689158],"study_design_scores_gemma":[0.000054284093,0.00032696125,0.02800287,0.0000827946,0.000055979595,0.00071826734,0.0009818177,0.92898786,0.025096025,0.010771118,0.0048564514,0.000065477805],"about_ca_topic_score_codex":0.0074210917,"about_ca_topic_score_gemma":0.009095062,"teacher_disagreement_score":0.0074210917,"about_ca_system_score_codex":0.00097476336,"about_ca_system_score_gemma":0.0017973653,"threshold_uncertainty_score":0.014755845},"labels":[],"label_agreement":null},{"id":"W1972275887","doi":"10.1109/saner.2015.7081831","title":"Detecting duplicate bug reports with software engineering domain knowledge","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Data deduplication; Context (archaeology); Domain (mathematical analysis); Software; Word (group theory); Software bug; Software engineering; Information retrieval; Artificial intelligence; Database; Programming language","score_opus":0.020912943414945238,"score_gpt":0.24827721961945887,"score_spread":0.22736427620451363,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972275887","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47467068,0.009268184,0.49024445,0.0009149506,0.00034838388,0.0019124694,0.005419278,0.011499947,0.005721663],"genre_scores_gemma":[0.47371542,0.0014139538,0.51064086,0.00030065564,0.00015810624,0.00093882706,0.009618325,0.00056773913,0.0026460537],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9844094,0.0028295396,0.0026536575,0.003339118,0.006359616,0.00040868888],"domain_scores_gemma":[0.8799061,0.059949763,0.023827324,0.013756389,0.021613527,0.0009469109],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009819553,0.0015034619,0.002094212,0.02899652,0.0013140511,0.0027971938,0.0028773574,0.0018902052,0.001368465],"category_scores_gemma":[0.09018364,0.0009238295,0.0012995407,0.013824465,0.00083420705,0.0043106955,0.0036851487,0.001164179,0.0009872095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065641786,0.0005716382,0.09980887,0.004334344,0.0006045668,0.003276968,0.0049312846,0.0063125617,0.028208638,0.0025370996,0.0073413975,0.8414162],"study_design_scores_gemma":[0.0004139603,0.0020408873,0.26008207,0.0025886784,0.0032778997,0.016975358,0.011808501,0.33642483,0.22612052,0.027245667,0.11213862,0.0008830214],"about_ca_topic_score_codex":0.003892877,"about_ca_topic_score_gemma":0.0062629064,"teacher_disagreement_score":0.02899652,"about_ca_system_score_codex":0.0010972228,"about_ca_system_score_gemma":0.0029303967,"threshold_uncertainty_score":0.05193138},"labels":[],"label_agreement":null},{"id":"W1972416872","doi":"10.1145/1808901.1808905","title":"Challenging cloning related problems with GPU-based algorithms","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Graphics processing unit; CUDA; Graphics; Longest common subsequence problem; General-purpose computing on graphics processing units; Parallel computing; Cloning (programming); False positive paradox; Matching (statistics); Computer graphics (images); Algorithm; Theoretical computer science; Artificial intelligence; Programming language; Mathematics","score_opus":0.011781084982681898,"score_gpt":0.2328380607661966,"score_spread":0.2210569757835147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972416872","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022021422,0.00052875275,0.9744188,0.00039177592,0.000066970366,0.000036424004,0.000022046352,0.000773033,0.0017407183],"genre_scores_gemma":[0.10637211,0.00038500095,0.89021444,0.00011819499,0.00006259385,0.00006552345,0.000091276284,0.00014568087,0.002545251],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985637,0.00040822694,0.00010195206,0.00021639043,0.0006057845,0.00010400245],"domain_scores_gemma":[0.99535626,0.0026111428,0.00032880742,0.0008288767,0.00077962835,0.00009526311],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018791346,0.0008059588,0.0011759547,0.0013808274,0.0011429753,0.002004883,0.0021021552,0.002323661,0.0029158886],"category_scores_gemma":[0.010208534,0.00060545217,0.00072806014,0.0025747335,0.0011333892,0.0024303324,0.0015998733,0.0016628797,0.0010561841],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032577832,0.00013963407,0.004367062,0.00033340562,0.00014145949,0.00052590057,0.0004395368,0.30176908,0.025973124,0.072591744,0.0070349863,0.58635825],"study_design_scores_gemma":[0.000040706418,0.00007142039,0.000454781,0.000018753391,0.000020205096,0.00038430854,0.00009019579,0.9376824,0.008724961,0.046771944,0.0057207304,0.000019617471],"about_ca_topic_score_codex":0.0037749938,"about_ca_topic_score_gemma":0.0036667574,"teacher_disagreement_score":0.0037749938,"about_ca_system_score_codex":0.0009858495,"about_ca_system_score_gemma":0.0010381587,"threshold_uncertainty_score":0.009937942},"labels":[],"label_agreement":null},{"id":"W1972466226","doi":"10.5555/2662708.2662714","title":"Scaling classical clone detection tools for ultra-large datasets: an exploratory study","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of Saskatchewan","funders":"","keywords":"Shuffling; Computer science; Scalability; Java; clone (Java method); Precision and recall; Cloning (programming); Data mining; Scaling; Machine learning; Artificial intelligence; Programming language; Database; Biology; Mathematics","score_opus":0.049494839322930875,"score_gpt":0.309086298293043,"score_spread":0.25959145897011215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972466226","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9399999,0.0010398387,0.046629984,0.00082224276,0.00014681417,0.0010675067,0.0029703248,0.0051414147,0.0021819838],"genre_scores_gemma":[0.83673155,0.00038580855,0.15314482,0.00030072228,0.00011179098,0.0008065779,0.0074325316,0.000483862,0.00060231163],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9804198,0.008309624,0.0019336155,0.0029327385,0.0055404757,0.0008637473],"domain_scores_gemma":[0.8724886,0.08492332,0.006467116,0.02406334,0.010066214,0.0019914452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01971687,0.0013021786,0.00084488944,0.0046716076,0.0013186635,0.0026027607,0.0028011347,0.0012808797,0.00074156706],"category_scores_gemma":[0.08053191,0.00045522998,0.0012319221,0.0050684796,0.0018026525,0.0055727414,0.0039050865,0.0024206731,0.00051906856],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002352144,0.0072692716,0.31871235,0.0026590652,0.001448106,0.001220604,0.0046081147,0.06384504,0.06805839,0.0069346516,0.04014085,0.4827514],"study_design_scores_gemma":[0.0007858631,0.004954597,0.184117,0.00033384768,0.0004253476,0.0018906527,0.0060755466,0.69426167,0.068895884,0.013851288,0.024117583,0.00029078836],"about_ca_topic_score_codex":0.002801332,"about_ca_topic_score_gemma":0.0033699148,"teacher_disagreement_score":0.01971687,"about_ca_system_score_codex":0.0013266306,"about_ca_system_score_gemma":0.0012538891,"threshold_uncertainty_score":0.104274035},"labels":[],"label_agreement":null},{"id":"W1972501133","doi":"10.5555/2486788.2487042","title":"Changeset based developer communication to detect software failures","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software; Software engineering; Software system; Software development; Software construction; Operating system","score_opus":0.036070879078120864,"score_gpt":0.2786316262526435,"score_spread":0.2425607471745226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972501133","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4838101,0.0009208218,0.4654598,0.001118998,0.00052043574,0.00096578815,0.003200428,0.0092238765,0.034779776],"genre_scores_gemma":[0.91967154,0.00017871754,0.07227546,0.00010087897,0.00007771947,0.00037319493,0.0009991843,0.00015662707,0.0061666667],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99712116,0.0008664546,0.00018348504,0.0005500356,0.0010563814,0.00022237819],"domain_scores_gemma":[0.98455435,0.0080522215,0.0023668904,0.0014449881,0.0029869806,0.00059454865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015799444,0.00049632613,0.0005221825,0.0037292514,0.000691676,0.0010357709,0.0007470919,0.0007997219,0.003285079],"category_scores_gemma":[0.015170337,0.0002571513,0.00028282922,0.0017435523,0.00033684055,0.0025652642,0.0012312451,0.0007326451,0.0011376977],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019502487,0.0011663006,0.18786253,0.00093157246,0.00037926107,0.0014001214,0.003350776,0.029643934,0.026668193,0.01844899,0.022375977,0.705822],"study_design_scores_gemma":[0.00014714827,0.0011495915,0.10035565,0.00015453174,0.0002803969,0.001509243,0.0028778487,0.79469025,0.038845737,0.019288609,0.040530242,0.00017078943],"about_ca_topic_score_codex":0.002590037,"about_ca_topic_score_gemma":0.00415576,"teacher_disagreement_score":0.0037292514,"about_ca_system_score_codex":0.0006283227,"about_ca_system_score_gemma":0.00069683074,"threshold_uncertainty_score":0.010989666},"labels":[],"label_agreement":null},{"id":"W1972786057","doi":"","title":"An Empirical Study of the Copy and Paste Behavior during Development","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Copying; Computer science; Eclipse; Context (archaeology); Code (set theory); Programming language; Source code","score_opus":0.05693587498825337,"score_gpt":0.34075599426623787,"score_spread":0.2838201192779845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972786057","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9972856,0.00019595069,0.0011191025,0.0000987088,0.0000041175217,0.000031532596,0.00026497213,0.000030170677,0.00096978445],"genre_scores_gemma":[0.9965868,0.00023761127,0.002013193,0.000055437988,0.000008552965,0.000054075117,0.00050231593,0.000028494362,0.000513457],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9891741,0.0044621127,0.0010917226,0.0016070317,0.0032201933,0.00044482655],"domain_scores_gemma":[0.8107898,0.12729667,0.034566455,0.009705788,0.015022208,0.0026190914],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0074504246,0.00032191313,0.00032071935,0.003127,0.0006303305,0.0015464644,0.0008021994,0.00073597045,0.0009286392],"category_scores_gemma":[0.077624194,0.0004918332,0.00023057662,0.0028525477,0.0009960718,0.0031038045,0.0010812532,0.0009238212,0.00044864626],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014062285,0.00025446757,0.9309011,0.00030261543,0.000089853595,0.00040627387,0.02330694,0.00027618874,0.0027884394,0.0003005707,0.0009768341,0.04025604],"study_design_scores_gemma":[0.000013117167,0.0003760914,0.968326,0.00013888707,0.00004303124,0.00116073,0.017727502,0.0027418686,0.0024542196,0.00033158532,0.006635411,0.00005158945],"about_ca_topic_score_codex":0.0017176945,"about_ca_topic_score_gemma":0.0030507194,"teacher_disagreement_score":0.9925496,"about_ca_system_score_codex":0.0005948234,"about_ca_system_score_gemma":0.00058858603,"threshold_uncertainty_score":0.039402127},"labels":[],"label_agreement":null},{"id":"W1972886276","doi":"10.1109/scam.2010.13","title":"Validating the Use of Topic Models for Software Evolution","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software evolution; Source code; Software maintenance; Software; Metric (unit); Software development; Task (project management); Code (set theory); Software system; Topic model; Software engineering; Data science; Software construction; Artificial intelligence; Programming language; Engineering; Systems engineering","score_opus":0.0754765022626967,"score_gpt":0.28401530902762945,"score_spread":0.20853880676493275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972886276","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6585702,0.0007987529,0.3265579,0.0014893643,0.00020867029,0.0015923542,0.002895928,0.0017841939,0.0061027375],"genre_scores_gemma":[0.87256795,0.0001798446,0.1221233,0.00015102004,0.000060849266,0.0012288783,0.0029357849,0.00025552712,0.0004969222],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9537335,0.03392285,0.0028282464,0.0049698516,0.003818524,0.0007270401],"domain_scores_gemma":[0.51264715,0.43533793,0.01252461,0.022392197,0.015328748,0.0017694675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.065652415,0.0013054867,0.0010733262,0.0055672578,0.0017395357,0.0047045527,0.0020792494,0.0024030542,0.0014457921],"category_scores_gemma":[0.290467,0.0008706403,0.0021177107,0.005007278,0.0019179125,0.008308609,0.002567565,0.00308201,0.0008606378],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031719664,0.0020049366,0.3827031,0.0020106945,0.001502095,0.00054273213,0.02620576,0.22915752,0.0077698478,0.037716553,0.00948177,0.297733],"study_design_scores_gemma":[0.00022830421,0.0007392296,0.03805954,0.00017679726,0.00012791705,0.00029909136,0.002206283,0.92936504,0.003622951,0.019813258,0.0052118204,0.00014968259],"about_ca_topic_score_codex":0.01134998,"about_ca_topic_score_gemma":0.009017676,"teacher_disagreement_score":0.065652415,"about_ca_system_score_codex":0.0035574278,"about_ca_system_score_gemma":0.0023342785,"threshold_uncertainty_score":0.3472073},"labels":[],"label_agreement":null},{"id":"W1973285032","doi":"10.1109/icdim.2013.6694005","title":"Maintenance support in open source software projects","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University; Western University","funders":"","keywords":"Backporting; Open source; Open source software; Software; Software maintenance; BitTorrent tracker; Software peer review; Software development; Software analytics","score_opus":0.02536941248220104,"score_gpt":0.27294649102533014,"score_spread":0.2475770785431291,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1973285032","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9982179,0.00013695199,0.000371047,0.000117095,0.0000032240964,0.000014520055,0.00012187687,0.000034208002,0.0009832818],"genre_scores_gemma":[0.99837315,0.000091218295,0.0006513541,0.000015586602,0.0000120771965,0.000028518118,0.00037263587,0.000008145666,0.0004473631],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99320316,0.0029319378,0.00063216925,0.00065129576,0.0021028705,0.0004785545],"domain_scores_gemma":[0.8558736,0.080583416,0.037953872,0.0065317713,0.012596828,0.0064605945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006692081,0.00018037352,0.00028745775,0.005105476,0.00081590074,0.0016276141,0.0008077691,0.000812885,0.0020212461],"category_scores_gemma":[0.0698463,0.0002247189,0.00021034587,0.0033451586,0.00039402305,0.0024523735,0.0020856597,0.0005173045,0.00044776805],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003356438,0.0007926004,0.8225896,0.00021508818,0.000031533058,0.0005855332,0.012014717,0.00092983054,0.0021475437,0.0007898485,0.0016417613,0.15792634],"study_design_scores_gemma":[0.000023326389,0.00030353488,0.9827702,0.00009506992,0.00002219069,0.00038578338,0.0049783634,0.0040419823,0.00097435096,0.0009508375,0.0054203086,0.000034044046],"about_ca_topic_score_codex":0.0020410449,"about_ca_topic_score_gemma":0.0024969096,"teacher_disagreement_score":0.006692081,"about_ca_system_score_codex":0.0008078954,"about_ca_system_score_gemma":0.0010131211,"threshold_uncertainty_score":0.03539151},"labels":[],"label_agreement":null},{"id":"W1974316556","doi":"10.1145/1082983.1083124","title":"A qualitative empirical evaluation of design decisions","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Outcome (game theory); Computer science; Management science; Decision analysis; Function (biology); Process (computing); Decision engineering; Business decision mapping; Operations research; Risk analysis (engineering); Decision support system; Artificial intelligence; Engineering; Mathematics; Mathematical economics","score_opus":0.18309279043444682,"score_gpt":0.4182378890722724,"score_spread":0.23514509863782557,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974316556","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5032864,0.0010864575,0.30838364,0.0062339185,0.0002865123,0.0066006924,0.0032380335,0.0002958205,0.17058857],"genre_scores_gemma":[0.90802735,0.00065998395,0.08009916,0.00048142867,0.000027269094,0.0035306204,0.0005073389,0.000056692443,0.0066101975],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9480069,0.041033674,0.0014347576,0.0012314497,0.007463274,0.0008299303],"domain_scores_gemma":[0.80319864,0.15871578,0.006173561,0.006120618,0.024661906,0.001129471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.041148927,0.00065737165,0.000454359,0.00472545,0.0028777542,0.0044009276,0.0012883286,0.0010904601,0.0074949116],"category_scores_gemma":[0.1376187,0.00041008685,0.00034838283,0.004151983,0.0062196325,0.00510226,0.0026465254,0.0013404774,0.00081006496],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008100075,0.0011286674,0.027895292,0.004557472,0.000079048674,0.00092303246,0.32459798,0.008107643,0.012716711,0.3566849,0.010498094,0.25200114],"study_design_scores_gemma":[0.0002477726,0.0018830412,0.019912567,0.004573694,0.000089922985,0.0008088115,0.60496515,0.02317074,0.024674635,0.15125771,0.16818382,0.00023205766],"about_ca_topic_score_codex":0.0017037934,"about_ca_topic_score_gemma":0.0027799031,"teacher_disagreement_score":0.041148927,"about_ca_system_score_codex":0.0054785963,"about_ca_system_score_gemma":0.0042921617,"threshold_uncertainty_score":0.21761894},"labels":[],"label_agreement":null},{"id":"W1974565163","doi":"10.1145/1985404.1985412","title":"Analyzing web service similarity using contextual clones","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Web service; clone (Java method); Leverage (statistics); Fragment (logic); World Wide Web; Context (archaeology); WS-I Basic Profile; Code (set theory); Programming language; Service (business); Information retrieval; Artificial intelligence; Web development; Web application security","score_opus":0.08931433237774358,"score_gpt":0.2932469354616403,"score_spread":0.20393260308389674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974565163","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89304376,0.00045189512,0.10327404,0.000070859256,0.0000132768955,0.00008897927,0.00018495595,0.00072778936,0.0021443984],"genre_scores_gemma":[0.9642044,0.00011156222,0.034950066,0.000021437172,0.00001673916,0.000058734633,0.00032877602,0.000051298044,0.00025706616],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975394,0.0004951719,0.00021736995,0.0005799258,0.0009945313,0.00017362768],"domain_scores_gemma":[0.988356,0.0053645675,0.0019695817,0.0016969467,0.0022564393,0.00035647597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014065914,0.00037756574,0.00074200705,0.006184341,0.0008447321,0.001477499,0.0006311839,0.00082672265,0.00073830714],"category_scores_gemma":[0.01884963,0.0002790156,0.000566493,0.0037637558,0.00081011443,0.002052491,0.0015428021,0.00043120523,0.00018713945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011517747,0.0003557902,0.42945588,0.00048696817,0.00037912032,0.0023139294,0.004495249,0.047469594,0.074449286,0.028050685,0.0016580224,0.40973374],"study_design_scores_gemma":[0.0000852972,0.0007815351,0.23426129,0.00013052208,0.0005525902,0.0059518944,0.003825,0.6416046,0.06675784,0.031253558,0.01457475,0.00022104481],"about_ca_topic_score_codex":0.0038563819,"about_ca_topic_score_gemma":0.0039411723,"teacher_disagreement_score":0.006184341,"about_ca_system_score_codex":0.00090487994,"about_ca_system_score_gemma":0.0006479446,"threshold_uncertainty_score":0.007667899},"labels":[],"label_agreement":null},{"id":"W1974583579","doi":"10.1109/csmr-wcre.2014.6747160","title":"Unification and refactoring of clones","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada; Concordia University","keywords":"Unification; Code refactoring; Programming language; Computer science; Software","score_opus":0.018184971254496442,"score_gpt":0.25639533125720626,"score_spread":0.23821036000270981,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974583579","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14478047,0.0009611219,0.8474592,0.00017817668,0.000071481176,0.00036209097,0.00009326211,0.0037532465,0.0023409931],"genre_scores_gemma":[0.35340923,0.00042921246,0.6402357,0.00014390692,0.000032327607,0.00019218455,0.00043237113,0.0006110738,0.004513969],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9950701,0.000900724,0.0005259461,0.001269082,0.0018618973,0.00037228162],"domain_scores_gemma":[0.9841993,0.0046460116,0.0020353007,0.0056691514,0.0031758083,0.0002744877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037483876,0.0009509818,0.0013953841,0.0029178439,0.00082408235,0.0009910662,0.0023301933,0.0014632228,0.0014570749],"category_scores_gemma":[0.015239073,0.0006442481,0.0015372268,0.0021816231,0.00080572226,0.0019692285,0.0022521964,0.0014828828,0.00058914477],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024422107,0.00019440202,0.008307212,0.0004769384,0.000112883885,0.0010409958,0.0013341912,0.023369277,0.14623998,0.011496753,0.0012326202,0.8059505],"study_design_scores_gemma":[0.0001689894,0.0009083707,0.01408492,0.00033742018,0.00052292144,0.004302436,0.000908107,0.5122604,0.39013505,0.028326552,0.04781647,0.00022844748],"about_ca_topic_score_codex":0.0024350423,"about_ca_topic_score_gemma":0.0025355106,"teacher_disagreement_score":0.0037483876,"about_ca_system_score_codex":0.0007881506,"about_ca_system_score_gemma":0.0014113628,"threshold_uncertainty_score":0.01982361},"labels":[],"label_agreement":null},{"id":"W1974601952","doi":"10.1109/icsm.2010.5609560","title":"Studying the impact of dependency network measures on software quality","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Betweenness centrality; Dependency (UML); Computer science; Eclipse; Software quality; Quality (philosophy); Centrality; Software; Network performance; Software bug; Reliability engineering; Data mining; Software development; Software engineering; Engineering; Computer network","score_opus":0.055649963916244875,"score_gpt":0.3516847744785293,"score_spread":0.2960348105622844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974601952","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97703815,0.00086370576,0.018315034,0.0003771797,0.00002194526,0.000020909984,0.00062062393,0.00008390637,0.0026585837],"genre_scores_gemma":[0.9968497,0.00014226687,0.0024676982,0.000009497622,0.000013138645,0.000011217905,0.00032594046,0.000023574148,0.00015690064],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9962435,0.0017692886,0.00020246985,0.0004912003,0.0010205525,0.00027300176],"domain_scores_gemma":[0.70252454,0.26335588,0.017565425,0.007269391,0.0075137415,0.0017710535],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.005430288,0.00065905973,0.000345583,0.0040200907,0.00040426967,0.00093951635,0.0005030569,0.0005085446,0.0013953083],"category_scores_gemma":[0.090935335,0.00022966256,0.0005212193,0.0041299337,0.0006285662,0.0032240988,0.00062657834,0.0011663218,0.00016580116],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035813975,0.00024590452,0.8053398,0.00028521914,0.0007442874,0.00020387523,0.000587124,0.09052035,0.0046703457,0.0047432124,0.0013629097,0.09093872],"study_design_scores_gemma":[0.000020056812,0.00039609856,0.7820657,0.00005262104,0.00026949076,0.0003110277,0.00036353077,0.20092267,0.0057788,0.008433398,0.0013234828,0.00006316422],"about_ca_topic_score_codex":0.004548298,"about_ca_topic_score_gemma":0.007050824,"teacher_disagreement_score":0.9945697,"about_ca_system_score_codex":0.0010747955,"about_ca_system_score_gemma":0.0005034366,"threshold_uncertainty_score":0.028718412},"labels":[],"label_agreement":null},{"id":"W1974712054","doi":"10.5555/2486788.2487004","title":"A study of variability spaces in open source software","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Maintainability; Computer science; Software; Software system; Space (punctuation); Software engineering; Work (physics); Programming language; Engineering; Operating system","score_opus":0.042680680918987794,"score_gpt":0.3053308972043857,"score_spread":0.2626502162853979,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974712054","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94639885,0.00028078997,0.050649937,0.0002680518,0.000007049748,0.00004020871,0.00009686941,0.00018287095,0.0020754815],"genre_scores_gemma":[0.98652065,0.000064438966,0.013077415,0.000017691318,0.000006198427,0.000023566017,0.00007177068,0.000047125766,0.00017122642],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9936939,0.0027991396,0.0003969242,0.0007977425,0.0020891181,0.00022308526],"domain_scores_gemma":[0.8453726,0.12355468,0.015731094,0.010024103,0.0039840895,0.0013334578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042104647,0.00028291202,0.0002717724,0.0022866207,0.0011483271,0.0024484396,0.0010892432,0.0009069766,0.000627565],"category_scores_gemma":[0.046960413,0.00046878232,0.00051463715,0.0030030021,0.0024051287,0.0050594555,0.0017343641,0.0018209751,0.00007018356],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008368639,0.0010179956,0.42459738,0.00070280896,0.00032181744,0.0028813807,0.06317177,0.098962635,0.038604394,0.1272429,0.0015943163,0.24006584],"study_design_scores_gemma":[0.00008544133,0.0013658839,0.2889815,0.00038270096,0.0002452493,0.0030800137,0.019764757,0.48057094,0.030059623,0.16075705,0.014362165,0.00034467486],"about_ca_topic_score_codex":0.002483505,"about_ca_topic_score_gemma":0.0026815096,"teacher_disagreement_score":0.0042104647,"about_ca_system_score_codex":0.0010209053,"about_ca_system_score_gemma":0.0011259124,"threshold_uncertainty_score":0.022267282},"labels":[],"label_agreement":null},{"id":"W1974789398","doi":"10.1109/compsac.2013.68","title":"The ReqWiki Approach for Collaborative Software Requirements Engineering with Integrated Text Analysis Support","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Requirements engineering; Software engineering; Software requirements; Software requirements specification; Requirements analysis; Requirement; Requirements elicitation; User requirements document; Software development; Software; World Wide Web; Software construction; Programming language","score_opus":0.018231863972953323,"score_gpt":0.2496935840366899,"score_spread":0.23146172006373658,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974789398","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027089692,0.00007641949,0.9672601,0.0003838206,0.00005479271,0.00063806167,0.0005753815,0.025077842,0.0032245473],"genre_scores_gemma":[0.017694129,0.00011363723,0.9718056,0.00022057637,0.000040625357,0.0008624363,0.0020059168,0.00249094,0.0047662677],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9862233,0.0053690504,0.0018867813,0.0021393667,0.0039879098,0.0003936346],"domain_scores_gemma":[0.96921176,0.0131992465,0.002667636,0.010075069,0.003847067,0.0009991736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012021631,0.0018749919,0.0011665868,0.005770789,0.0015338785,0.0043791505,0.0053886883,0.0020644218,0.0076188906],"category_scores_gemma":[0.03184149,0.0017682533,0.0026459454,0.0027329975,0.0016553608,0.008719942,0.008704198,0.0039220466,0.0062640104],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094041444,0.0015957599,0.0019084227,0.0038869309,0.0006018359,0.0022874998,0.009224403,0.016603889,0.10783892,0.117063604,0.050607685,0.6874406],"study_design_scores_gemma":[0.0005049347,0.00043533675,0.0018312978,0.0004904227,0.00026380384,0.0028921657,0.0014175904,0.18645705,0.080798335,0.13299792,0.59129685,0.0006143452],"about_ca_topic_score_codex":0.0021289587,"about_ca_topic_score_gemma":0.0044900193,"teacher_disagreement_score":0.012021631,"about_ca_system_score_codex":0.0009780993,"about_ca_system_score_gemma":0.0037818323,"threshold_uncertainty_score":0.063577235},"labels":[],"label_agreement":null},{"id":"W1974908123","doi":"10.5555/2819009.2819054","title":"Code repurposing as an assessment tool","year":2015,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sheridan College","funders":"","keywords":"Repurposing; Computer science; Code (set theory); Software engineering; Source code; Software; Code review; KPI-driven code analysis; Programming language; Software development; Static program analysis; Engineering; Set (abstract data type)","score_opus":0.06233157398028721,"score_gpt":0.36394204481772896,"score_spread":0.30161047083744175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974908123","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33033267,0.0010137708,0.5599827,0.003565915,0.0006956789,0.0032369217,0.00046182182,0.02736936,0.073341206],"genre_scores_gemma":[0.61098105,0.0005304709,0.36245987,0.00072240003,0.00019635736,0.0012993655,0.0008317015,0.0024579775,0.02052076],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9773013,0.011191586,0.0010820513,0.0013397003,0.008224851,0.00086042576],"domain_scores_gemma":[0.9210273,0.03957439,0.0050894176,0.007630346,0.022938732,0.0037398182],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012365981,0.0012426581,0.0008492877,0.0059316596,0.0010455596,0.004337313,0.002061008,0.0012226077,0.0079936115],"category_scores_gemma":[0.073105715,0.000474,0.0005156005,0.0022660787,0.0009432713,0.00513452,0.0041831196,0.002425326,0.0047858576],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042755355,0.0014562095,0.0075482735,0.00080841215,0.00003166142,0.0010221797,0.0114928065,0.002857838,0.033490356,0.0057544294,0.020953508,0.9141569],"study_design_scores_gemma":[0.0005817434,0.009039286,0.06963216,0.003297716,0.0003157547,0.014776191,0.017741216,0.15282497,0.20216319,0.043220088,0.48515168,0.0012560693],"about_ca_topic_score_codex":0.00043316197,"about_ca_topic_score_gemma":0.0006516932,"teacher_disagreement_score":0.987634,"about_ca_system_score_codex":0.00078331673,"about_ca_system_score_gemma":0.0014184491,"threshold_uncertainty_score":0.065398335},"labels":[],"label_agreement":null},{"id":"W1975143584","doi":"10.1109/wcre.2007.27","title":"Examining the Effects of Global Data Usage on Software Maintainability","year":2007,"lang":"en","type":"article","venue":"Proceedings - Working Conference on Reverse Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Maintainability; Computer science; Software maintenance; Source code; Software; Software evolution; Variable (mathematics); Source lines of code; Task (project management); Software metric; Software development; Software engineering; Software construction; Operating system; Systems engineering; Engineering","score_opus":0.046603457429799246,"score_gpt":0.2815427781236253,"score_spread":0.2349393206938261,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1975143584","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.998055,0.000099114666,0.0010972061,0.000041581196,0.0000030521371,0.000011035166,0.00008247253,0.000028417553,0.0005821355],"genre_scores_gemma":[0.9991437,0.00003645864,0.00055353483,0.000008763042,0.000006119404,0.000010371822,0.00011186739,0.000015129531,0.00011400113],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9896909,0.0047121174,0.0008677732,0.001305274,0.0028582073,0.00056584727],"domain_scores_gemma":[0.6066442,0.3097914,0.05021264,0.016722532,0.01379536,0.002833816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0093831355,0.0006294544,0.00051522,0.0027384316,0.00040645938,0.0015081987,0.0006436132,0.00059099623,0.0009749018],"category_scores_gemma":[0.10779241,0.00032936898,0.0007494114,0.0031038963,0.0008973757,0.0027082544,0.001446009,0.0012071691,0.00018616415],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033432606,0.00019054138,0.9720977,0.000068261616,0.00029291882,0.00011653005,0.0008803019,0.0023906631,0.0024014907,0.0001418584,0.000073703195,0.021011785],"study_design_scores_gemma":[0.000005354424,0.00066063204,0.99388784,0.0000112886155,0.0000922767,0.00011207152,0.00030498992,0.0030835734,0.0015388085,0.00011062248,0.00017622246,0.00001639803],"about_ca_topic_score_codex":0.003235677,"about_ca_topic_score_gemma":0.0038985482,"teacher_disagreement_score":0.0093831355,"about_ca_system_score_codex":0.0006227387,"about_ca_system_score_gemma":0.0005191873,"threshold_uncertainty_score":0.04962337},"labels":[],"label_agreement":null},{"id":"W1975197452","doi":"10.1145/2259051.2259058","title":"Collection disjointness analysis","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Dataflow; Java; Computer science; Static analysis; Programming language; Inference; Theoretical computer science; Artificial intelligence","score_opus":0.017875349469819463,"score_gpt":0.27402345159985975,"score_spread":0.2561481021300403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1975197452","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07319149,0.00071110035,0.90255356,0.000426518,0.00011826698,0.0004165416,0.001924644,0.009494055,0.011163865],"genre_scores_gemma":[0.5023367,0.0006203236,0.47989637,0.00049008755,0.00024325917,0.00068688614,0.0050807507,0.0040305676,0.0066150404],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9857819,0.0013216458,0.0010829602,0.0025724447,0.0078228945,0.0014180366],"domain_scores_gemma":[0.9750292,0.009658133,0.0027400719,0.0062137,0.00583358,0.00052534725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040836283,0.0013516092,0.0014903566,0.008093727,0.0030418872,0.0031828042,0.0030244312,0.0010566022,0.0077111167],"category_scores_gemma":[0.022104373,0.0010563049,0.00348428,0.004700025,0.0027477618,0.008568922,0.005567403,0.0023412872,0.0012267564],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008766883,0.00039901963,0.06214148,0.0015986927,0.0005620477,0.0024698798,0.0028041506,0.032878846,0.08105346,0.33470237,0.017434625,0.4630787],"study_design_scores_gemma":[0.00009668485,0.00043448003,0.02380741,0.0006059897,0.00089097145,0.003383206,0.0011572472,0.19176105,0.245363,0.41050255,0.12159004,0.00040734693],"about_ca_topic_score_codex":0.0038267677,"about_ca_topic_score_gemma":0.0030565783,"teacher_disagreement_score":0.008093727,"about_ca_system_score_codex":0.0016631451,"about_ca_system_score_gemma":0.0037598065,"threshold_uncertainty_score":0.025796235},"labels":[],"label_agreement":null},{"id":"W1975318342","doi":"10.1145/2512207","title":"Degree-of-knowledge","year":2014,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; International Business Machines Corporation","keywords":"Codebase; Computer science; Robustness (evolution); Source code; Code (set theory); Code review; Point (geometry); Software; Software engineering; Software development; Static program analysis; Programming language; Set (abstract data type)","score_opus":0.11510190202852359,"score_gpt":0.3348329567299539,"score_spread":0.21973105470143028,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1975318342","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056185424,0.0005019475,0.9034237,0.0021681276,0.00008193305,0.0002288064,0.0009312434,0.0004868903,0.03599189],"genre_scores_gemma":[0.8683514,0.0004428699,0.12172078,0.00023448032,0.00013034153,0.00035541487,0.0007136973,0.00012605479,0.007924987],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9890665,0.0030448872,0.0007625224,0.0032458145,0.0030716117,0.00080874766],"domain_scores_gemma":[0.9524715,0.02929797,0.0035716859,0.008569116,0.0044339527,0.0016556348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006494386,0.0010347078,0.001257253,0.0044274637,0.0014351876,0.0054447385,0.0031945945,0.0026362124,0.008682928],"category_scores_gemma":[0.05491863,0.0007869877,0.0017768416,0.0037241227,0.0043905103,0.014377599,0.003775098,0.0023293323,0.001555151],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020481006,0.00019895373,0.021840807,0.00041765615,0.00022879655,0.00034748908,0.002368189,0.07776704,0.0016529579,0.7810506,0.0049548205,0.10896787],"study_design_scores_gemma":[0.00003744187,0.00010592884,0.0058447276,0.000100478675,0.00010887673,0.000627529,0.00048046547,0.19349031,0.0011053799,0.7793666,0.018662585,0.00006974406],"about_ca_topic_score_codex":0.004637455,"about_ca_topic_score_gemma":0.0033858228,"teacher_disagreement_score":0.008682928,"about_ca_system_score_codex":0.0031620504,"about_ca_system_score_gemma":0.0016820568,"threshold_uncertainty_score":0.034346044},"labels":[],"label_agreement":null},{"id":"W1975917686","doi":"10.1109/iri.2012.6302984","title":"Reusing and converting code clones to aspects - An algorithmic approach","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cistel Technology (Canada); Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Redundant code; Source code; Code (set theory); Unreachable code; Dead code; KPI-driven code analysis; Cloning (programming); Programming language; Code generation; Code reuse; Static program analysis; Reuse; Constant-weight code; Systematic code; clone (Java method); Algorithm; Operating system; Software; Software development; Set (abstract data type); Linear code; Decoding methods; Engineering; Code rate; Key (lock)","score_opus":0.03565660403092491,"score_gpt":0.2902317827957717,"score_spread":0.2545751787648468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1975917686","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006625383,0.0000818321,0.9908347,0.000100424615,0.000012134222,0.00014794935,0.000015930234,0.0010010018,0.0011805949],"genre_scores_gemma":[0.04154163,0.00015946754,0.9564857,0.00007299606,0.000010863561,0.00009517648,0.00012149737,0.00029513994,0.0012175034],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948887,0.001118395,0.00058893254,0.00096321653,0.0021741495,0.00026649138],"domain_scores_gemma":[0.9888826,0.003854671,0.00087491143,0.004483449,0.0017561004,0.0001482461],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036958945,0.0009585383,0.0006347513,0.002442529,0.0009960448,0.003618072,0.0022893024,0.0012066098,0.0021910144],"category_scores_gemma":[0.01596169,0.0009818358,0.001632953,0.0014954132,0.0024553067,0.0033492425,0.0023964788,0.0018709565,0.0009329653],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013330142,0.00035078873,0.007797898,0.0006835472,0.00018771375,0.00054285134,0.0016348384,0.049934078,0.07033676,0.14887215,0.0018614719,0.7176646],"study_design_scores_gemma":[0.00014368162,0.0006086117,0.004525735,0.0006336475,0.00036152638,0.003638145,0.0011369237,0.45579985,0.24021818,0.16987835,0.12284239,0.00021301802],"about_ca_topic_score_codex":0.0017396276,"about_ca_topic_score_gemma":0.0020346718,"teacher_disagreement_score":0.0036958945,"about_ca_system_score_codex":0.000886417,"about_ca_system_score_gemma":0.0020659065,"threshold_uncertainty_score":0.019545972},"labels":[],"label_agreement":null},{"id":"W1975985860","doi":"10.1109/vissoft.2013.6650524","title":"Visualizing software dynamicities with heat maps","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Visualization; Program comprehension; Software visualization; Software; Comprehension; Software evolution; Class (philosophy); Software engineering; Data visualization; Software development; Programming language; Human–computer interaction; Software system; Software construction; Data mining; Artificial intelligence","score_opus":0.0103015550478777,"score_gpt":0.23738924441775816,"score_spread":0.22708768936988047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1975985860","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022860076,0.00038864164,0.959202,0.00057298737,0.00006893477,0.00008204972,0.00078226725,0.01139496,0.0046481327],"genre_scores_gemma":[0.4121562,0.0009948283,0.5809757,0.00014920424,0.000107790285,0.00040147,0.00095416565,0.0018371418,0.0024235926],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995524,0.00019564125,0.000023321621,0.00006612814,0.00011795972,0.00004441302],"domain_scores_gemma":[0.99680954,0.002296133,0.00019472592,0.000331866,0.00023454787,0.0001330755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009887146,0.0012267908,0.00047003088,0.003990026,0.00062310114,0.0027797127,0.00090519426,0.0008832403,0.009087995],"category_scores_gemma":[0.006021786,0.00043149432,0.00086112175,0.002115505,0.0009087342,0.0037905227,0.002445148,0.0015325352,0.00068096473],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096589053,0.00028005228,0.008197885,0.0017322114,0.0003216089,0.0013676404,0.018023435,0.09928369,0.07371666,0.29766572,0.03189067,0.4665545],"study_design_scores_gemma":[0.0001574454,0.00024654233,0.007862498,0.00042239623,0.00016304998,0.0008688811,0.0027593386,0.4469134,0.042736687,0.37404093,0.12350045,0.0003284619],"about_ca_topic_score_codex":0.0017979866,"about_ca_topic_score_gemma":0.0016437866,"teacher_disagreement_score":0.009087995,"about_ca_system_score_codex":0.00045465035,"about_ca_system_score_gemma":0.000509998,"threshold_uncertainty_score":0.030402422},"labels":[],"label_agreement":null},{"id":"W1976148265","doi":"10.1109/icsm.2012.6405248","title":"Vasco: A visual approach to explore object churn in framework-intensive applications","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Visualization; Task (project management); Scalability; Object (grammar); Visual analytics; Human–computer interaction; Task analysis; Artificial intelligence; Database; Engineering; Systems engineering","score_opus":0.054186893091562695,"score_gpt":0.3327363668041033,"score_spread":0.2785494737125406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1976148265","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051980782,0.00059511134,0.90660405,0.0009604203,0.00009147205,0.00028469844,0.0012212354,0.03302908,0.005233168],"genre_scores_gemma":[0.27838618,0.0006522338,0.7131056,0.00022737567,0.0000701485,0.0004649732,0.0009616945,0.0034401806,0.002691586],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995919,0.00013956023,0.000023557437,0.00006587319,0.00012435114,0.000054685155],"domain_scores_gemma":[0.997306,0.0016284725,0.00021332927,0.00030433302,0.00030675382,0.00024112589],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012994969,0.0015079597,0.00049439183,0.0033088275,0.0008703152,0.0022054645,0.0015442359,0.0011157506,0.004789618],"category_scores_gemma":[0.004604154,0.0006864177,0.0007249342,0.0012436315,0.0006819122,0.002145447,0.0032908001,0.0020483395,0.00047591986],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014977314,0.0005147887,0.014293492,0.0020132943,0.00038039678,0.0027659019,0.023266265,0.07611859,0.15339662,0.04432895,0.067886055,0.61353785],"study_design_scores_gemma":[0.00033221394,0.0004795867,0.020187957,0.0006818424,0.00022188654,0.001986127,0.0037126837,0.67056966,0.05883396,0.06733117,0.17520234,0.00046052615],"about_ca_topic_score_codex":0.0038635049,"about_ca_topic_score_gemma":0.0058304397,"teacher_disagreement_score":0.004789618,"about_ca_system_score_codex":0.00050404615,"about_ca_system_score_gemma":0.0009096689,"threshold_uncertainty_score":0.016022861},"labels":[],"label_agreement":null},{"id":"W1977136723","doi":"10.1109/wcre.2010.46","title":"Software Process Recovery: Recovering Process from Artifacts","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software development; Software engineering; Process (computing); Team software process; Goal-Driven Software Development Process; Software development process; Personal software process; Instrumentation (computer programming); Software project management; Software; Software system; Software construction; Programming language","score_opus":0.012706940147016033,"score_gpt":0.2676927547580854,"score_spread":0.2549858146110694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1977136723","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013559533,0.000120983255,0.9820331,0.00021287742,0.000019969482,0.00015512276,0.00016481488,0.0031397627,0.0005938744],"genre_scores_gemma":[0.19462867,0.00041614266,0.7996172,0.000107468135,0.000039190298,0.0003184926,0.0017588004,0.000963698,0.002150385],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99184716,0.0022528544,0.0005731182,0.0017326474,0.0031671827,0.00042694478],"domain_scores_gemma":[0.9682519,0.010117497,0.0051260022,0.011567764,0.0045482432,0.00038861262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005183009,0.0015034691,0.001070674,0.007163121,0.001218372,0.0028714864,0.0026532589,0.0017006326,0.0015004863],"category_scores_gemma":[0.035162434,0.0007429538,0.0018529244,0.005083738,0.0021718554,0.0065847947,0.003153957,0.0025428205,0.0014026613],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026723088,0.00034375876,0.013662146,0.00086015783,0.00015071353,0.0005507856,0.0043381415,0.042954344,0.024331154,0.029637514,0.0055925976,0.87731135],"study_design_scores_gemma":[0.00008822355,0.0005077733,0.025997926,0.0005620645,0.00028594138,0.0016638563,0.0037269667,0.6810092,0.10738484,0.1166391,0.061807163,0.00032699254],"about_ca_topic_score_codex":0.0041508498,"about_ca_topic_score_gemma":0.0027826035,"teacher_disagreement_score":0.007163121,"about_ca_system_score_codex":0.000882405,"about_ca_system_score_gemma":0.0028038013,"threshold_uncertainty_score":0.027410686},"labels":[],"label_agreement":null},{"id":"W1977490305","doi":"10.1109/ase.2009.62","title":"Improving API Usage through Automatic Detection of Redundant Code","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Application programming interface; Java; Software; Code (set theory); Software quality; Code smell; Software engineering; Software development; Programming language; Operating system","score_opus":0.01680323653613325,"score_gpt":0.2686442627403735,"score_spread":0.25184102620424026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1977490305","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.59454095,0.0005635093,0.3807548,0.00037568828,0.000060138333,0.00048205917,0.00055317424,0.020335997,0.002333644],"genre_scores_gemma":[0.6812611,0.00017907441,0.31493005,0.000085531945,0.000028921935,0.00024368038,0.0009822679,0.0009805065,0.0013088958],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9897702,0.002699836,0.0010522164,0.0021688577,0.0038754889,0.00043345356],"domain_scores_gemma":[0.9165982,0.03251856,0.018457152,0.015113287,0.016273879,0.0010388942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004376704,0.0012417779,0.0011192429,0.005870679,0.0007543259,0.001540383,0.0023913477,0.0011328992,0.0007150811],"category_scores_gemma":[0.05189678,0.0009703466,0.0006763097,0.00269824,0.000971703,0.002621077,0.0017281093,0.0015125099,0.0007177756],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048008614,0.00072340906,0.14027002,0.00093559816,0.00020453504,0.00088160107,0.0026882123,0.016019529,0.1174979,0.0020386542,0.0042577772,0.7140027],"study_design_scores_gemma":[0.00012856795,0.00078018534,0.09747647,0.00019526725,0.00034099948,0.0028265375,0.0008033169,0.65399426,0.2260794,0.006174061,0.010899949,0.0003010085],"about_ca_topic_score_codex":0.0019226439,"about_ca_topic_score_gemma":0.0030152302,"teacher_disagreement_score":0.005870679,"about_ca_system_score_codex":0.0005443144,"about_ca_system_score_gemma":0.0017419395,"threshold_uncertainty_score":0.02314651},"labels":[],"label_agreement":null},{"id":"W1977500793","doi":"10.1109/wcre.2012.58","title":"Exploring How to Develop Transformations and Tools for Automated Umplification","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Programming language; Programmer; Code (set theory); USable; Unified Modeling Language; Code generation; Software engineering; Process (computing); Java; Representation (politics); Software; Operating system","score_opus":0.19654602823107398,"score_gpt":0.31661127993874466,"score_spread":0.12006525170767068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1977500793","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032378114,0.000100916404,0.9892565,0.0003167033,0.000029441848,0.0001323948,0.000052304535,0.0050722454,0.0018017184],"genre_scores_gemma":[0.033231597,0.0002487042,0.96329933,0.0001316768,0.000013356742,0.00012564436,0.00035634814,0.0013107447,0.0012826221],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99414885,0.002472688,0.00047398199,0.0008744129,0.0016606519,0.0003694232],"domain_scores_gemma":[0.98220164,0.009043758,0.001084938,0.00534261,0.002090767,0.0002362907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062168157,0.0015984944,0.000744546,0.0023914077,0.0010947787,0.0034647388,0.0027723706,0.0016666006,0.005076582],"category_scores_gemma":[0.029096333,0.0011979332,0.0021639266,0.0013679699,0.002061256,0.008093918,0.0038650576,0.003052041,0.0034524326],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017839526,0.00055116037,0.003493233,0.0013443776,0.00012277799,0.001226184,0.0024347599,0.040761452,0.038438868,0.14268611,0.008382832,0.76037985],"study_design_scores_gemma":[0.00014872495,0.00036061893,0.0016981045,0.0010364399,0.0001835769,0.0025844425,0.0019965502,0.33808967,0.16741171,0.2339826,0.25226176,0.00024577585],"about_ca_topic_score_codex":0.0014261038,"about_ca_topic_score_gemma":0.001655858,"teacher_disagreement_score":0.0062168157,"about_ca_system_score_codex":0.0008441174,"about_ca_system_score_gemma":0.002552494,"threshold_uncertainty_score":0.03287804},"labels":[],"label_agreement":null},{"id":"W1977601746","doi":"10.1109/compsac.2013.66","title":"Semantic-Enabled Clone Detection","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Social Semantic Web; Semantic Web; Semantics (computer science); Semantic Web Stack; Source code; Semantic analytics; Semantic matching; Information retrieval; Domain (mathematical analysis); Semantic computing; clone (Java method); Semantic grid; Semantic search; Focus (optics); World Wide Web; Matching (statistics); Programming language","score_opus":0.010287995319958781,"score_gpt":0.22708942547996236,"score_spread":0.21680143016000358,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1977601746","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016860716,0.0003885398,0.9737579,0.00044585494,0.000117720454,0.0001409795,0.00024753387,0.0035909854,0.0044498052],"genre_scores_gemma":[0.24718983,0.00047877984,0.74736595,0.00039902088,0.000082679944,0.00013515918,0.0009146131,0.000499448,0.002934432],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99515754,0.0008461822,0.00037706786,0.00084479456,0.0023746914,0.00039975508],"domain_scores_gemma":[0.9905724,0.0028791332,0.0008823638,0.0028129132,0.0025829826,0.00027022394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030841364,0.0008449865,0.0010363488,0.0050566387,0.0011117815,0.0032696023,0.0020723261,0.0021465356,0.0025069118],"category_scores_gemma":[0.016286504,0.0005425422,0.001443681,0.0034233874,0.0024928248,0.0069055585,0.0044237217,0.0014856743,0.0011607629],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039171777,0.00032281113,0.01682197,0.0009037507,0.00018250076,0.0022505082,0.0018788419,0.023911929,0.059368774,0.26698917,0.00798129,0.61899674],"study_design_scores_gemma":[0.000068465946,0.00029227027,0.005074974,0.00028912345,0.00024530358,0.004496435,0.0007354727,0.3923496,0.15300055,0.33882165,0.1043931,0.00023309108],"about_ca_topic_score_codex":0.0021495025,"about_ca_topic_score_gemma":0.0022977171,"teacher_disagreement_score":0.0050566387,"about_ca_system_score_codex":0.0010044518,"about_ca_system_score_gemma":0.00202476,"threshold_uncertainty_score":0.016310692},"labels":[],"label_agreement":null},{"id":"W1978248929","doi":"10.1109/iwsm.mensura.2014.44","title":"COSMIC Approximate Sizing Using a Fuzzy Logic Approach: A Quantitative Case Study with Industry Data","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Sizing; Fuzzy logic; Computer science; Software; Process (computing); Industrial engineering; Reliability engineering; Data mining; Artificial intelligence; Engineering","score_opus":0.18777196495159304,"score_gpt":0.3605434953177218,"score_spread":0.17277153036612874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1978248929","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8986537,0.000216544,0.090980716,0.00033574225,0.000012971865,0.00024990752,0.0005755171,0.00015224676,0.0088227615],"genre_scores_gemma":[0.95177114,0.000069160385,0.04724702,0.000018256913,0.0000042891634,0.00008114837,0.00017398312,0.000017481047,0.0006174682],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968348,0.0017852831,0.00013942066,0.00021974798,0.0009197386,0.00010104556],"domain_scores_gemma":[0.9821977,0.0146848485,0.0005329386,0.0008593526,0.0016064461,0.000118605414],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005348365,0.000456488,0.0004801459,0.0020452014,0.0007441006,0.0010504419,0.0010968106,0.0010597733,0.0017799989],"category_scores_gemma":[0.013669559,0.00025714762,0.00044448944,0.0027378376,0.0009111871,0.0012644142,0.0005035074,0.00073528336,0.00013103742],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010173102,0.0010434311,0.031265337,0.00068711076,0.0001084632,0.0016210203,0.0034174263,0.77953595,0.0080440445,0.04050478,0.002021586,0.13073365],"study_design_scores_gemma":[0.00009033885,0.0005996591,0.009577464,0.000057167206,0.00003447234,0.00038281895,0.0018254338,0.9708811,0.0050780177,0.007891411,0.003531055,0.000051055085],"about_ca_topic_score_codex":0.008542707,"about_ca_topic_score_gemma":0.008536756,"teacher_disagreement_score":0.008542707,"about_ca_system_score_codex":0.0019113407,"about_ca_system_score_gemma":0.000656153,"threshold_uncertainty_score":0.028285205},"labels":[],"label_agreement":null},{"id":"W1978832807","doi":"10.5555/2820518.2820554","title":"An empirical study of end-user programmers in the computer music community","year":2015,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Metadata; Software; World Wide Web; Software development; Population; Source code; Computer programming; Musical; End user; Empirical research; Software engineering; Programming language","score_opus":0.0705591971718184,"score_gpt":0.3386892397372468,"score_spread":0.2681300425654284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1978832807","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99903584,0.000056969984,0.00017284851,0.00012349687,0.000001863625,0.000026038397,0.000020053883,0.00000412603,0.00055863237],"genre_scores_gemma":[0.9982899,0.00019009669,0.00065901823,0.0001431852,0.000008824586,0.00006451961,0.00008751175,0.000008732681,0.00054825173],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99141407,0.004314546,0.00058627425,0.00079991634,0.0020147692,0.0008704363],"domain_scores_gemma":[0.8854943,0.06875053,0.020179577,0.0027631135,0.013030615,0.009781898],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007805523,0.00030166196,0.00041571129,0.005157612,0.0028630118,0.0029269713,0.0013957402,0.0010472479,0.0020262415],"category_scores_gemma":[0.061599486,0.00050290034,0.00016023291,0.0045376197,0.0021331364,0.0044298917,0.002997827,0.0015478492,0.0003479936],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014358315,0.0013157371,0.669469,0.0002171801,0.00002702413,0.0009925697,0.29476866,0.0000627695,0.0012321023,0.0006285951,0.0010412364,0.030101495],"study_design_scores_gemma":[0.000031408184,0.00067154365,0.3825352,0.00015250104,0.00001904247,0.0012177125,0.6085589,0.0008065532,0.0005769348,0.00038440136,0.0050013233,0.000044497472],"about_ca_topic_score_codex":0.0051866216,"about_ca_topic_score_gemma":0.010546165,"teacher_disagreement_score":0.007805523,"about_ca_system_score_codex":0.0009886025,"about_ca_system_score_gemma":0.0020211258,"threshold_uncertainty_score":0.04128009},"labels":[],"label_agreement":null},{"id":"W1978859404","doi":"10.1109/ase.2013.6693087","title":"Personalized defect prediction","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":243,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software bug; Java; Commit; Eclipse; Source lines of code; Python (programming language); Software; Code (set theory); Software evolution; Machine learning; Predictive modelling; Kernel (algebra); Linux kernel; Data mining; Artificial intelligence; Programming language; Operating system; Software development; Database; Software construction","score_opus":0.013632446480468364,"score_gpt":0.2366302322925787,"score_spread":0.22299778581211036,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1978859404","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4657866,0.0016988928,0.48408774,0.0011131341,0.00023652034,0.00041086398,0.011385324,0.02903387,0.0062470045],"genre_scores_gemma":[0.8789325,0.0004931695,0.10119373,0.00022618953,0.00011567845,0.00019555741,0.012571684,0.00060045737,0.005671102],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998365,0.00017141437,0.00009122947,0.00060046074,0.0006161558,0.00015572546],"domain_scores_gemma":[0.99140114,0.003220923,0.0013065537,0.0017984228,0.001986161,0.00028672913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010965614,0.0013790624,0.0010612832,0.0041553676,0.00026585843,0.00082294334,0.001544872,0.0011319019,0.0019065567],"category_scores_gemma":[0.007566814,0.0004255445,0.0009114549,0.0019644417,0.0002183657,0.0016658676,0.00082559773,0.0010001947,0.0015063593],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038493393,0.000668414,0.27929923,0.00027197553,0.00027861036,0.0005209224,0.0001928818,0.08230295,0.007927264,0.00092314393,0.025175517,0.60205424],"study_design_scores_gemma":[0.000037251393,0.0003454543,0.06212856,0.000042632757,0.00016529788,0.0007553306,0.00011878537,0.91408056,0.011650763,0.0037645581,0.00683693,0.000073882584],"about_ca_topic_score_codex":0.005104979,"about_ca_topic_score_gemma":0.007841356,"teacher_disagreement_score":0.005104979,"about_ca_system_score_codex":0.00045848417,"about_ca_system_score_gemma":0.0005389292,"threshold_uncertainty_score":0.010150552},"labels":[],"label_agreement":null},{"id":"W1978917996","doi":"10.1145/2020390.2020406","title":"Customization support for CBR-based defect prediction","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Case-based reasoning; Weighting; Data mining; k-nearest neighbors algorithm; Similarity (geometry); Personalization; Set (abstract data type); Artificial intelligence; Function (biology); Machine learning; Selection (genetic algorithm)","score_opus":0.036519441192453,"score_gpt":0.25525581944004294,"score_spread":0.21873637824758996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1978917996","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057765596,0.00041367172,0.8258899,0.0004956799,0.00007654741,0.0007034833,0.0009138044,0.10584416,0.007897179],"genre_scores_gemma":[0.5665596,0.00049949327,0.42038506,0.00039893072,0.00015477187,0.0005877983,0.0035757707,0.0030762397,0.004762348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99660504,0.000546914,0.00047753484,0.00082977605,0.0012988518,0.00024188541],"domain_scores_gemma":[0.9791461,0.00796659,0.0018492853,0.008107031,0.002470156,0.00046091148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035587247,0.001138199,0.00092719943,0.001804121,0.00033269625,0.0024664316,0.005277818,0.0013392747,0.0055739647],"category_scores_gemma":[0.02138689,0.0007103877,0.0009638279,0.00189256,0.0006705661,0.0031945035,0.0019969817,0.0019031479,0.002876203],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011227676,0.0012014367,0.013286834,0.000747032,0.00026217563,0.0015058762,0.0010436054,0.08596601,0.038559224,0.019717785,0.02294188,0.8136453],"study_design_scores_gemma":[0.00030715988,0.00021293151,0.0063361125,0.00025098174,0.00019883603,0.0016670087,0.00012900584,0.87357354,0.05124909,0.02707588,0.038809575,0.00018988053],"about_ca_topic_score_codex":0.002554261,"about_ca_topic_score_gemma":0.0015626997,"teacher_disagreement_score":0.0055739647,"about_ca_system_score_codex":0.000862621,"about_ca_system_score_gemma":0.00058415503,"threshold_uncertainty_score":0.018820584},"labels":[],"label_agreement":null},{"id":"W1979217904","doi":"10.1109/ictta.2008.4530015","title":"Functional Equivalence between Radial Basis Function Neural Networks and Fuzzy Analogy in Software Cost Estimation","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Analogy; Equivalence (formal languages); Computer science; Radial basis function network; Fuzzy logic; Artificial neural network; Artificial intelligence; Fuzzy set; Radial basis function; Mathematics; Machine learning; Discrete mathematics","score_opus":0.03762357696809028,"score_gpt":0.25840987002925625,"score_spread":0.22078629306116598,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979217904","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03008123,0.00069768337,0.9599963,0.00057786447,0.00006489244,0.000016806915,0.000031332595,0.000054568776,0.008479283],"genre_scores_gemma":[0.86214757,0.00073997275,0.13317038,0.00023794781,0.00017845788,0.000090903515,0.00008890179,0.000052728454,0.0032931503],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99630725,0.0018764598,0.00018099486,0.0005026916,0.0009467277,0.00018595226],"domain_scores_gemma":[0.9955819,0.0027429606,0.00039144183,0.00040522596,0.00077634875,0.00010211624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046965666,0.00056591746,0.0007883703,0.0009637723,0.0004568225,0.0013779378,0.0012416082,0.0016183656,0.002183697],"category_scores_gemma":[0.015167481,0.0002685878,0.0009504764,0.00073987956,0.0022678378,0.004279324,0.0019511033,0.0017016336,0.00029797663],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000447612,0.000038673952,0.0006275448,0.0000768013,0.00004035797,0.00009561148,0.00016723957,0.08978871,0.0015120232,0.8738865,0.0004463862,0.033275276],"study_design_scores_gemma":[0.000006429996,0.00006199394,0.0004960998,0.000018318577,0.000014603589,0.00008163814,0.00003167111,0.44599625,0.000599155,0.5514603,0.0012128701,0.00002067181],"about_ca_topic_score_codex":0.0020914725,"about_ca_topic_score_gemma":0.00089749455,"teacher_disagreement_score":0.0046965666,"about_ca_system_score_codex":0.0009906207,"about_ca_system_score_gemma":0.00063966983,"threshold_uncertainty_score":0.02483809},"labels":[],"label_agreement":null},{"id":"W1979296382","doi":"10.5555/2486788.2486979","title":"Using mutation analysis for a model-clone detector comparison framework","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Computer science; Mutation; Precision and recall; Data mining; Artificial intelligence; Machine learning; Biology; Genetics; Gene","score_opus":0.0743848086925239,"score_gpt":0.349684522917882,"score_spread":0.2752997142253581,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979296382","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00680742,0.000041301762,0.9906448,0.00011251561,0.0000135288055,0.000090029396,0.000022044356,0.0018096905,0.000458569],"genre_scores_gemma":[0.15037425,0.00004188099,0.84852105,0.00010810169,0.000014122554,0.00012788332,0.00009777038,0.00035355106,0.0003613986],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98360544,0.006293729,0.0014332263,0.0018769033,0.006023902,0.0007668548],"domain_scores_gemma":[0.97559434,0.012011967,0.003297072,0.0033490364,0.0052865962,0.0004608976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017226446,0.0014582819,0.0014987013,0.0056707496,0.0013392249,0.0046354746,0.0039658737,0.0023564447,0.0017795839],"category_scores_gemma":[0.04151288,0.0008135302,0.0026777198,0.0022857687,0.0027397298,0.0049378234,0.0038247074,0.002721058,0.00038441134],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000508105,0.00077115407,0.019161094,0.00047183552,0.00050338596,0.0009856956,0.0010650809,0.3163645,0.04219313,0.25577945,0.0037015178,0.358495],"study_design_scores_gemma":[0.00003720824,0.00019445644,0.0009184106,0.0000894251,0.000103208964,0.00027503027,0.0001223872,0.9284775,0.0269818,0.03790758,0.0048161494,0.00007686345],"about_ca_topic_score_codex":0.0060270596,"about_ca_topic_score_gemma":0.0046934066,"teacher_disagreement_score":0.017226446,"about_ca_system_score_codex":0.0033264838,"about_ca_system_score_gemma":0.0050575263,"threshold_uncertainty_score":0.091103196},"labels":[],"label_agreement":null},{"id":"W1979464590","doi":"10.1109/compsacw.2012.67","title":"SE-EQUAM - An Evolvable Quality Metamodel","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Evolvability; Metamodeling; Computer science; Software engineering; Reusability; Reuse; Software mining; Maintainability; Quality (philosophy); Software; Data science; Software system; Engineering; Software construction; Programming language","score_opus":0.1241667663361443,"score_gpt":0.38159522533272855,"score_spread":0.25742845899658423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979464590","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012868404,0.00017564154,0.97622967,0.0005927377,0.00006219274,0.00021048542,0.00033688525,0.0015392741,0.00798461],"genre_scores_gemma":[0.1766044,0.0003591021,0.81601834,0.00028527176,0.000033725784,0.00044238815,0.0010387303,0.00033856995,0.0048795007],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9963421,0.000981399,0.0005640749,0.00044666702,0.0014372638,0.00022842933],"domain_scores_gemma":[0.99443406,0.001302811,0.0006634761,0.0016715832,0.0016934638,0.00023461143],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063533816,0.000665324,0.0004918974,0.0021762892,0.00082725973,0.0031314045,0.0018604712,0.0014550304,0.0015491555],"category_scores_gemma":[0.010208645,0.00057317916,0.0019657996,0.0009731409,0.0014386695,0.0039016434,0.0026949893,0.0022109349,0.00053772464],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017454792,0.00030747813,0.00868854,0.00067272125,0.00021396566,0.0005639734,0.0019376182,0.106077895,0.018960385,0.68174547,0.004946534,0.17571102],"study_design_scores_gemma":[0.00007875016,0.0003051004,0.002282837,0.0007769087,0.00023187851,0.0009044582,0.0006915045,0.47975543,0.024947776,0.30735004,0.18252423,0.00015114617],"about_ca_topic_score_codex":0.0034384977,"about_ca_topic_score_gemma":0.0057705236,"teacher_disagreement_score":0.0063533816,"about_ca_system_score_codex":0.0015391565,"about_ca_system_score_gemma":0.0029433558,"threshold_uncertainty_score":0.03360027},"labels":[],"label_agreement":null},{"id":"W1979506693","doi":"10.1087/20130307","title":"Republication of conference papers in journals?","year":2013,"lang":"en","type":"article","venue":"Learned Publishing","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Similarity (geometry); Field (mathematics); Discretion; Quarter (Canadian coin); Library science; Political science; Law; History; Mathematics; Artificial intelligence","score_opus":0.0422389926299143,"score_gpt":0.28357975549560566,"score_spread":0.24134076286569137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979506693","genre_codex":"empirical","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51991355,0.025550142,0.017699746,0.10926227,0.026967825,0.0025668961,0.008772139,0.00453466,0.28473273],"genre_scores_gemma":[0.93805194,0.0026791648,0.007768265,0.01117713,0.0032239014,0.0010566245,0.0026302943,0.0014742849,0.03193844],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.7526685,0.104688495,0.04881079,0.014146104,0.069850795,0.0098352665],"domain_scores_gemma":[0.18511562,0.33585256,0.199771,0.13002463,0.127804,0.021432294],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.18529591,0.0005341458,0.0012457863,0.015879298,0.0063719545,0.021161212,0.004498858,0.0029077386,0.038118508],"category_scores_gemma":[0.64420027,0.00074565195,0.0011697994,0.023448896,0.0039611054,0.015989643,0.008302384,0.0039268765,0.013128881],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031673284,0.0007175177,0.21007456,0.0046024076,0.0006111541,0.0013924289,0.026117418,0.0005917749,0.0043339487,0.075934015,0.2383565,0.43410102],"study_design_scores_gemma":[0.0003669632,0.0007664763,0.25512108,0.004105508,0.00026474887,0.0029008647,0.021514704,0.0019791455,0.009918071,0.025808418,0.67692554,0.00032851828],"about_ca_topic_score_codex":0.0011659414,"about_ca_topic_score_gemma":0.000832987,"teacher_disagreement_score":0.9788388,"about_ca_system_score_codex":0.00891434,"about_ca_system_score_gemma":0.008668488,"threshold_uncertainty_score":0.97995013},"labels":[],"label_agreement":null},{"id":"W1979613104","doi":"10.1145/2601248.2601249","title":"Quality control practice based on design artifacts categories","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Component (thermodynamics); Software engineering; Software quality; Software development; Warrant; Quality (philosophy); Artifact (error); Component-based software engineering; Software evolution; Control (management); Empirical research; Process (computing); Software design; Software; Software construction; Programming language; Artificial intelligence","score_opus":0.040984090277561544,"score_gpt":0.3177373840452483,"score_spread":0.27675329376768676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979613104","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51938796,0.0041213445,0.36019287,0.014673016,0.00018870925,0.0012107791,0.000062000305,0.0006153553,0.099547975],"genre_scores_gemma":[0.93742275,0.0005357248,0.059102,0.000411625,0.000022656504,0.00037431502,0.000058268313,0.00006424287,0.0020083361],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8485328,0.06905603,0.01121336,0.009401461,0.05890588,0.0028905573],"domain_scores_gemma":[0.7420629,0.14121373,0.028818052,0.037626345,0.046704877,0.003574078],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06648381,0.00073916186,0.00061411806,0.010356689,0.0032743067,0.011594373,0.003453343,0.0022223387,0.0009851467],"category_scores_gemma":[0.15149637,0.0007523664,0.00057026243,0.0046333186,0.016784064,0.008910335,0.007007175,0.0026108676,0.0001951368],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013585338,0.00044079564,0.072745755,0.0013646275,0.00009838286,0.00051132316,0.17736693,0.0033018275,0.005549516,0.37131983,0.0025124718,0.36465275],"study_design_scores_gemma":[0.00019669025,0.0011332507,0.07490448,0.0071266354,0.0001989376,0.0023733021,0.15893947,0.028470537,0.013134069,0.5403595,0.17284656,0.00031656507],"about_ca_topic_score_codex":0.003692888,"about_ca_topic_score_gemma":0.0026733174,"teacher_disagreement_score":0.06648381,"about_ca_system_score_codex":0.010593974,"about_ca_system_score_gemma":0.012482926,"threshold_uncertainty_score":0.35160422},"labels":[],"label_agreement":null},{"id":"W1979786964","doi":"10.1109/scam.2010.22","title":"Estimating the Optimal Number of Latent Concepts in Source Code Analysis","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Latent Dirichlet allocation; Computer science; Source code; Latent variable; Code (set theory); Topic model; Latent semantic analysis; Software; Natural language processing; Data mining; Artificial intelligence; Programming language","score_opus":0.01564829991397315,"score_gpt":0.31615772867474007,"score_spread":0.3005094287607669,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979786964","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2352991,0.00086797075,0.76006347,0.0011585421,0.00004699319,0.00024620086,0.0005658463,0.0010768472,0.0006750573],"genre_scores_gemma":[0.5985008,0.0004779832,0.3967328,0.00033143724,0.000108443324,0.0006588862,0.002135999,0.00034730864,0.0007063224],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9880553,0.0073481817,0.0005870724,0.0024495132,0.00088770833,0.00067228766],"domain_scores_gemma":[0.9058748,0.082054764,0.0024775665,0.005592824,0.0026157151,0.0013843217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016998133,0.0014448803,0.0028055082,0.0034448009,0.0023270668,0.0037349516,0.0038022646,0.004708469,0.0014129127],"category_scores_gemma":[0.094115995,0.002229939,0.00273088,0.0030685412,0.0040786127,0.010581556,0.0046631526,0.00632225,0.0009075566],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003571478,0.0015700841,0.03426526,0.00085560983,0.0005396176,0.00036503788,0.0038749177,0.62297285,0.01231813,0.048965942,0.006878625,0.26382247],"study_design_scores_gemma":[0.0001644449,0.00009290792,0.0018226868,0.0000692534,0.00006193743,0.00007637085,0.0002829067,0.9375305,0.0022081619,0.057010457,0.00062148634,0.00005880291],"about_ca_topic_score_codex":0.007541207,"about_ca_topic_score_gemma":0.008112636,"teacher_disagreement_score":0.016998133,"about_ca_system_score_codex":0.0028789106,"about_ca_system_score_gemma":0.0030272924,"threshold_uncertainty_score":0.089895785},"labels":[],"label_agreement":null},{"id":"W1979810153","doi":"10.1145/1370905.1370913","title":"Security metrics for source code structures","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"H2020 European Research Council; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software security assurance; Source code; Software metric; Software; Code (set theory); KPI-driven code analysis; Software development; Static program analysis; Software quality; Software engineering; Computer security; Programming language; Information security; Security service; Set (abstract data type)","score_opus":0.033479366803898736,"score_gpt":0.2829550111991585,"score_spread":0.24947564439525977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979810153","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07547215,0.005281337,0.89117014,0.0012629628,0.00027337708,0.0011575609,0.00629279,0.006360598,0.012729146],"genre_scores_gemma":[0.35905656,0.0013372519,0.6289426,0.000114053495,0.00012097083,0.0011938857,0.0064033275,0.00081923476,0.0020121909],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97353554,0.006142608,0.0038999037,0.0016507296,0.014044897,0.00072633725],"domain_scores_gemma":[0.90026945,0.040066183,0.021258283,0.012622996,0.024392108,0.0013910326],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011464679,0.0018983957,0.00084841053,0.017630393,0.001179706,0.0032099094,0.0013616807,0.0015706209,0.0026215606],"category_scores_gemma":[0.10960047,0.0005874604,0.0012710014,0.011673865,0.0015487418,0.006671175,0.0023999903,0.0015073661,0.0006593896],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029724982,0.00028532336,0.07147649,0.0024834222,0.00045449715,0.000491943,0.0016445522,0.118300706,0.022149675,0.18756542,0.017944518,0.57690614],"study_design_scores_gemma":[0.000104992825,0.0012552624,0.07533898,0.0014490503,0.00036589263,0.0023401813,0.0008094525,0.5403002,0.035913862,0.25299025,0.08873287,0.00039894128],"about_ca_topic_score_codex":0.0024105888,"about_ca_topic_score_gemma":0.0026361009,"teacher_disagreement_score":0.98853534,"about_ca_system_score_codex":0.002579071,"about_ca_system_score_gemma":0.0021248164,"threshold_uncertainty_score":0.060631752},"labels":[],"label_agreement":null},{"id":"W1979829546","doi":"10.4018/jdm.2008010101","title":"Dimensions of UML Diagram Use","year":2008,"lang":"en","type":"article","venue":"Journal of Database Management","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland; University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Unified Modeling Language; Applications of UML; UML tool; Computer science; Class diagram; Use Case Diagram; Communication diagram; Software engineering; Systems Modeling Language; Software; Programming language","score_opus":0.045295651888089844,"score_gpt":0.2773889981208894,"score_spread":0.23209334623279956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979829546","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5277675,0.0086889295,0.17156334,0.012011661,0.00035608778,0.0012030835,0.0021832823,0.0011181535,0.27510798],"genre_scores_gemma":[0.9262448,0.0018097628,0.06672887,0.00046140168,0.00011915218,0.0006135772,0.0008780793,0.00016451893,0.0029798034],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9197799,0.040943477,0.009127799,0.0026372622,0.025746854,0.0017646487],"domain_scores_gemma":[0.8458104,0.10222856,0.020381622,0.009838548,0.018358901,0.0033818553],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023318157,0.00069466914,0.00035643013,0.011244338,0.0019030382,0.007963603,0.0008527137,0.0016110406,0.002474261],"category_scores_gemma":[0.08971489,0.00054514356,0.00067664107,0.008799191,0.004491765,0.008875574,0.005611106,0.0017740155,0.0006095126],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025160625,0.00028192403,0.1598877,0.0012960698,0.00023052252,0.0005487782,0.17439929,0.0021834907,0.007840458,0.2851963,0.0099902395,0.3578938],"study_design_scores_gemma":[0.00006177775,0.00038626822,0.13069776,0.002481197,0.00013350231,0.0053257686,0.078588635,0.0064047966,0.0036247044,0.21618706,0.55578226,0.0003263222],"about_ca_topic_score_codex":0.001915217,"about_ca_topic_score_gemma":0.0015333323,"teacher_disagreement_score":0.023318157,"about_ca_system_score_codex":0.0028014784,"about_ca_system_score_gemma":0.0022234228,"threshold_uncertainty_score":0.123319685},"labels":[],"label_agreement":null},{"id":"W1979838076","doi":"10.1016/j.future.2010.09.007","title":"An approach based on citation analysis to support effective handling of regulatory compliance","year":2010,"lang":"en","type":"article","venue":"Future Generation Computer Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Interdependence; Citation; Legislature; Scope (computer science); Compliance (psychology); Set (abstract data type); Risk analysis (engineering); Law; Operations research; Business; World Wide Web; Political science","score_opus":0.023556306430406907,"score_gpt":0.275662024699318,"score_spread":0.2521057182689111,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979838076","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013608225,0.00058469886,0.9540849,0.002931159,0.00057707925,0.00066912395,0.0011294403,0.0065356577,0.01987978],"genre_scores_gemma":[0.15486716,0.00058222347,0.83406794,0.00036441372,0.00061900634,0.00059562706,0.001288299,0.00043671837,0.007178645],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9783943,0.0066998624,0.0022430078,0.0022881823,0.009787805,0.0005868876],"domain_scores_gemma":[0.9038778,0.04306729,0.009668786,0.012893652,0.029390445,0.0011019752],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.015096906,0.0012542234,0.0017686994,0.030039025,0.003256801,0.008125984,0.0032912958,0.0035122093,0.00752049],"category_scores_gemma":[0.08597827,0.00063102634,0.0017195417,0.01704447,0.0016713246,0.011749671,0.003287733,0.0026769154,0.0036529587],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019900303,0.0006478257,0.0109741865,0.00071534346,0.0003685937,0.00038685097,0.0012622284,0.014596076,0.009554992,0.18747047,0.027283879,0.74654055],"study_design_scores_gemma":[0.00016786682,0.00034556031,0.010596028,0.0005553975,0.000735573,0.0009571601,0.001271851,0.4596359,0.027830714,0.38748667,0.10984212,0.00057526276],"about_ca_topic_score_codex":0.0050775358,"about_ca_topic_score_gemma":0.006956594,"teacher_disagreement_score":0.969961,"about_ca_system_score_codex":0.0017654048,"about_ca_system_score_gemma":0.008173853,"threshold_uncertainty_score":0.07984102},"labels":[],"label_agreement":null},{"id":"W1979991500","doi":"10.1007/s10664-012-9224-x","title":"Configuring latent Dirichlet allocation based feature location","year":2012,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"U.S. Department of Education; National Science Foundation","keywords":"Latent Dirichlet allocation; Computer science; Feature (linguistics); Source code; Heuristics; Context (archaeology); Java; Artificial intelligence; Topic model; Code (set theory); Measure (data warehouse); Data mining; Natural language processing; Information retrieval; Programming language","score_opus":0.024902158345490696,"score_gpt":0.2743756658225636,"score_spread":0.24947350747707292,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979991500","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02798516,0.00018843859,0.9646391,0.00031779503,0.00012259265,0.00008263373,0.00032344516,0.00535415,0.0009866538],"genre_scores_gemma":[0.4894229,0.00012699878,0.503398,0.00036012544,0.00014429285,0.00040108917,0.0021226562,0.0009832985,0.0030406285],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952041,0.0023403782,0.00027782223,0.0012115211,0.0005941085,0.0003720744],"domain_scores_gemma":[0.98939437,0.0063823764,0.00028595288,0.0023862347,0.0012302066,0.0003208873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046919524,0.0011722925,0.0019940569,0.001771715,0.0011988453,0.002288966,0.0035654048,0.0029960263,0.0059461007],"category_scores_gemma":[0.026178025,0.0011174115,0.0015627481,0.0020366644,0.0011304445,0.0043413946,0.0044274866,0.0025514,0.00408932],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028060828,0.0006807047,0.011383732,0.00021700388,0.00036026502,0.0003402931,0.00068906037,0.17105553,0.0198756,0.01983182,0.016937558,0.7558224],"study_design_scores_gemma":[0.00009839551,0.000056941408,0.00061625656,0.000011835741,0.00005144429,0.00008170229,0.000094595445,0.9660085,0.006740224,0.024545614,0.0016655826,0.000028918985],"about_ca_topic_score_codex":0.005126292,"about_ca_topic_score_gemma":0.0069454857,"teacher_disagreement_score":0.0059461007,"about_ca_system_score_codex":0.0012473048,"about_ca_system_score_gemma":0.0017349288,"threshold_uncertainty_score":0.024813771},"labels":[],"label_agreement":null},{"id":"W1980130600","doi":"10.1007/s10664-014-9329-5","title":"Special issue on program comprehension","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Program comprehension; Computer science; Comprehension; Data science; Psychology; Programming language; Software","score_opus":0.021811196545500226,"score_gpt":0.29366010915216767,"score_spread":0.27184891260666744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1980130600","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011580094,0.029325766,0.0035373315,0.0758656,0.77429914,0.00016459535,0.0012807695,0.000635164,0.11373373],"genre_scores_gemma":[0.0042727776,0.0167839,0.00096330454,0.010639122,0.6883703,0.00018668386,0.0022598014,0.0007136253,0.27581063],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980794,0.0003017076,0.00016433683,0.00034874014,0.00086265634,0.00024317861],"domain_scores_gemma":[0.9901154,0.0028528455,0.00058399583,0.0010536802,0.0031822075,0.0022118727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026461454,0.0021313396,0.0025027762,0.006820924,0.0020791267,0.0068121497,0.0023601782,0.004863947,0.19952083],"category_scores_gemma":[0.0090997,0.0006864965,0.0013021879,0.004038854,0.0012781651,0.0058516664,0.0046474207,0.0045236517,0.07241065],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028572023,0.000045407713,0.00012330378,0.00017769364,0.00000795802,0.00003791085,0.000023462999,0.000043125478,0.00018575165,0.0016720854,0.9651267,0.03252801],"study_design_scores_gemma":[0.000016166101,0.000048420552,0.0010803014,0.00028672622,0.000015027721,0.00012879187,0.000045849512,0.0001321158,0.00016046806,0.004967036,0.9931076,0.0000114949125],"about_ca_topic_score_codex":0.00093493186,"about_ca_topic_score_gemma":0.0023611442,"teacher_disagreement_score":0.19952083,"about_ca_system_score_codex":0.001855172,"about_ca_system_score_gemma":0.0027264266,"threshold_uncertainty_score":0.66746366},"labels":[],"label_agreement":null},{"id":"W1980905535","doi":"10.1109/ccece.2012.6335063","title":"Performance norms: An approach to rework reduction in software development","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Rework; Computer science; Task (project management); Software; Software development; Software engineering; Engineering management; Risk analysis (engineering); Engineering; Systems engineering; Business","score_opus":0.030763131703031295,"score_gpt":0.25977087507893176,"score_spread":0.22900774337590046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1980905535","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13521764,0.0016552417,0.7918361,0.0074886773,0.00039103828,0.0010035193,0.00020651381,0.00057816185,0.061623063],"genre_scores_gemma":[0.8024075,0.0005221797,0.1911106,0.0004897242,0.00025438733,0.0016040833,0.00012650742,0.00015295275,0.0033321262],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9190571,0.055434685,0.0026160018,0.0059238584,0.015769154,0.0011993205],"domain_scores_gemma":[0.83770627,0.103777446,0.02814692,0.010821804,0.017037814,0.0025097425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04760749,0.0018470656,0.0012728524,0.007823273,0.003160031,0.006208087,0.0044045104,0.0019588338,0.0019852028],"category_scores_gemma":[0.1429094,0.0006044057,0.0011234366,0.0045497427,0.012845939,0.0076973843,0.0055805845,0.003372943,0.00034781735],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003601563,0.0012216124,0.084402144,0.0006892727,0.00029431973,0.00029155775,0.048601694,0.014620954,0.002187014,0.4572161,0.003356456,0.3867588],"study_design_scores_gemma":[0.00014382399,0.0014760416,0.108598635,0.0011977317,0.00034475894,0.00056807185,0.042906284,0.14170606,0.006257429,0.65666515,0.03960875,0.000527323],"about_ca_topic_score_codex":0.007080071,"about_ca_topic_score_gemma":0.0066225384,"teacher_disagreement_score":0.04760749,"about_ca_system_score_codex":0.008460422,"about_ca_system_score_gemma":0.008380212,"threshold_uncertainty_score":0.2517755},"labels":[],"label_agreement":null},{"id":"W1981075560","doi":"10.1007/s10664-013-9274-8","title":"Studying the relationship between logging characteristics and the code quality of platform software","year":2013,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software; Product metric; Process (computing); Relation (database); Source lines of code; Logging; Software quality; Code (set theory); Database; Quality (philosophy); Product (mathematics); Software development; Software engineering; Data mining; Operating system; Programming language; Set (abstract data type)","score_opus":0.1071977212306796,"score_gpt":0.3313187722431654,"score_spread":0.2241210510124858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1981075560","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99929667,0.000029966688,0.00038427988,0.00002842717,9.790556e-7,0.0000026136277,0.000023450122,0.0000057768993,0.00022774888],"genre_scores_gemma":[0.99960035,0.000018083621,0.00020681089,0.000004415531,0.0000021831586,0.0000018234712,0.000052806656,0.0000041307358,0.00010935617],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99800843,0.0006725447,0.00019854162,0.00030302146,0.0005433173,0.00027406792],"domain_scores_gemma":[0.8040397,0.14255577,0.037680328,0.005471352,0.007526554,0.002726274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032256271,0.00026962408,0.00016723873,0.0017237697,0.0004061348,0.0015063495,0.0005910607,0.00068484363,0.002097047],"category_scores_gemma":[0.067163475,0.000351052,0.00037905216,0.0026370964,0.00074996264,0.0023992267,0.00074343185,0.0013404033,0.00028787873],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012693941,0.00017590798,0.99280506,0.000015329546,0.00006362877,0.000043674267,0.00023187899,0.0009770009,0.00082951004,0.00024090018,0.000039454837,0.0044508306],"study_design_scores_gemma":[0.000008034489,0.00024231337,0.98972356,0.000007778283,0.000053543165,0.0000952793,0.0005756542,0.007289979,0.0013805238,0.00046945593,0.00014192717,0.000011943666],"about_ca_topic_score_codex":0.0048947893,"about_ca_topic_score_gemma":0.008909918,"teacher_disagreement_score":0.0048947893,"about_ca_system_score_codex":0.00074850395,"about_ca_system_score_gemma":0.001101173,"threshold_uncertainty_score":0.017058909},"labels":[],"label_agreement":null},{"id":"W1981126395","doi":"10.1145/1081706.1081761","title":"Anchoring and adjustment in software estimation","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Anchoring; Respondent; Estimator; Estimation; Software; Computer science; Econometrics; Statistics; Psychology; Mathematics; Social psychology; Economics","score_opus":0.01446841534856634,"score_gpt":0.26556350620437263,"score_spread":0.2510950908558063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1981126395","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44127834,0.005619433,0.4899287,0.006265798,0.0005854302,0.0006729347,0.00014684064,0.0005357331,0.054966673],"genre_scores_gemma":[0.94815445,0.00089897995,0.048770037,0.0006014281,0.0002245772,0.00024710444,0.00005435557,0.00012375558,0.0009252408],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8371921,0.11882341,0.008371041,0.010178055,0.023219217,0.0022161538],"domain_scores_gemma":[0.31725848,0.5802409,0.05054099,0.03398753,0.016556263,0.0014157731],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08661341,0.0011042367,0.0009968836,0.0031863328,0.002090953,0.0067628934,0.0016406992,0.0035571558,0.0043767705],"category_scores_gemma":[0.60788256,0.0014022447,0.0013532348,0.0042024674,0.009860045,0.011109187,0.005674851,0.005156381,0.00077112345],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002184127,0.00046747818,0.14788924,0.0016804228,0.0012306372,0.0007054928,0.056670994,0.028511604,0.00847187,0.27729556,0.0036725674,0.47122002],"study_design_scores_gemma":[0.00043414952,0.0012574208,0.16923487,0.0015923689,0.00085251493,0.0011401712,0.0112961,0.067840844,0.014310449,0.6993961,0.031712394,0.00093270536],"about_ca_topic_score_codex":0.0036634267,"about_ca_topic_score_gemma":0.00195756,"teacher_disagreement_score":0.08661341,"about_ca_system_score_codex":0.0026363013,"about_ca_system_score_gemma":0.0018495903,"threshold_uncertainty_score":0.45806098},"labels":[],"label_agreement":null},{"id":"W1981223477","doi":"10.1142/s0218194008003763","title":"RELIABILITY MODEL FOR COMPONENT-BASED SYSTEMS IN COSMIC (A CASE STUDY)","year":2008,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; École de Technologie Supérieure; Concordia University","funders":"","keywords":"Component (thermodynamics); Reliability engineering; Reliability (semiconductor); Computer science; Markov chain; Context (archaeology); Markov model; Markov process; Probabilistic logic; Dependability; Component-based software engineering; Software system; Distributed computing; Systems engineering; Engineering; Software; Machine learning; Artificial intelligence","score_opus":0.027262164607725074,"score_gpt":0.27827164325388154,"score_spread":0.25100947864615647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1981223477","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22431237,0.0014503306,0.7346078,0.0010898396,0.000086236185,0.00030728805,0.0009880104,0.00092006184,0.03623801],"genre_scores_gemma":[0.94994444,0.00046631022,0.040393937,0.00006750538,0.000049610848,0.0002723039,0.00046205698,0.000075136166,0.008268681],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99912125,0.00032921147,0.00005350991,0.00013925819,0.0002557513,0.000101039615],"domain_scores_gemma":[0.9987538,0.0006543215,0.00014458016,0.000093216666,0.00030956574,0.000044569748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014772303,0.00077641616,0.00057717453,0.0011812449,0.0006389987,0.0013013057,0.001203197,0.0018021117,0.0024646104],"category_scores_gemma":[0.002633905,0.00028687785,0.00097195763,0.0008862726,0.00083395507,0.0008247442,0.0006963249,0.000775527,0.0005245026],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000065463224,0.000037160178,0.0018070857,0.00006830684,0.000025666957,0.0007094651,0.0003014538,0.93045485,0.0015492545,0.056871835,0.0013113396,0.006798085],"study_design_scores_gemma":[0.000009466671,0.000042659194,0.00055147643,0.000012456385,0.000018767101,0.00015500249,0.00005171553,0.98750985,0.00027852305,0.009336065,0.002024251,0.000009839885],"about_ca_topic_score_codex":0.013649394,"about_ca_topic_score_gemma":0.006813144,"teacher_disagreement_score":0.013649394,"about_ca_system_score_codex":0.0015181581,"about_ca_system_score_gemma":0.00094084657,"threshold_uncertainty_score":0.027139902},"labels":[],"label_agreement":null},{"id":"W1981598140","doi":"","title":"Ecosystems in GitHub and a Method for Ecosystem Identification using Reference Coupling","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ecosystem; Identification (biology); Computer science; Software; Coupling (piping); Isolation (microbiology); Data science; Environmental resource management; Software engineering; Ecology; Environmental science; Engineering","score_opus":0.06542503582081172,"score_gpt":0.3427650770691965,"score_spread":0.2773400412483848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1981598140","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07887998,0.00027435189,0.9102598,0.00017587596,0.000053678843,0.00023148462,0.00052694924,0.005667386,0.003930514],"genre_scores_gemma":[0.30679077,0.00013972599,0.68870366,0.00006420647,0.000020375945,0.0003580849,0.0011302576,0.0008204422,0.0019724914],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99346304,0.0026126427,0.000740771,0.0011331304,0.001721184,0.0003292623],"domain_scores_gemma":[0.9822953,0.0070731454,0.0021075117,0.0042481814,0.0035951803,0.00068075897],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.005323477,0.0008126124,0.00061687187,0.01038707,0.0011202195,0.003016515,0.0012483911,0.0012955256,0.002512694],"category_scores_gemma":[0.029692344,0.0006852095,0.0011829096,0.0068019726,0.0006714867,0.0032260166,0.0030898286,0.0012326292,0.0008628536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044051776,0.00037226366,0.15251827,0.0010094391,0.0005580985,0.0019132044,0.008392357,0.040222507,0.027510768,0.0656687,0.008892154,0.69250166],"study_design_scores_gemma":[0.00005763205,0.00023966956,0.07223143,0.00042656425,0.0002574231,0.0020958509,0.0027349323,0.7801508,0.025467781,0.057725113,0.058362503,0.00025031046],"about_ca_topic_score_codex":0.007012988,"about_ca_topic_score_gemma":0.009880494,"teacher_disagreement_score":0.98961294,"about_ca_system_score_codex":0.0010702463,"about_ca_system_score_gemma":0.0015297534,"threshold_uncertainty_score":0.028153598},"labels":[],"label_agreement":null},{"id":"W1981621595","doi":"10.1007/s10664-010-9132-x","title":"Testing the theory of relative defect proneness for closed-source software","year":2010,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Children's Hospital of Eastern Ontario","funders":"","keywords":"Open source software; Software; Open source; Computer science; Software engineering; Data science; Programming language","score_opus":0.0349592160170909,"score_gpt":0.28196653910387043,"score_spread":0.24700732308677953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1981621595","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98313934,0.00006072074,0.013686033,0.00033224517,0.000021416785,0.000034052933,0.00010809808,0.000053048414,0.002565046],"genre_scores_gemma":[0.9979159,0.000017439092,0.0017329851,0.000034271277,0.000015541356,0.000028933107,0.0001139423,0.000012618103,0.00012827308],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97975576,0.011212095,0.0010048938,0.003900984,0.0033391921,0.0007871195],"domain_scores_gemma":[0.34976456,0.60609597,0.019355701,0.01574102,0.00684461,0.002198206],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03089593,0.0011432688,0.00071107683,0.0025275634,0.0007718414,0.0019869825,0.0028777295,0.0022870654,0.006203123],"category_scores_gemma":[0.26443797,0.00044491273,0.0013975863,0.0017017145,0.0046493914,0.006115169,0.0023542396,0.0024349678,0.00045014624],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003945046,0.003972724,0.7239221,0.0006484842,0.0023822344,0.0005690422,0.0038523658,0.058184784,0.0074898987,0.11041311,0.0018536262,0.082766585],"study_design_scores_gemma":[0.001019995,0.008686923,0.4147912,0.00017311831,0.0007627332,0.0010283676,0.0033891576,0.38493633,0.008474749,0.17514142,0.0014425487,0.00015340424],"about_ca_topic_score_codex":0.0012065943,"about_ca_topic_score_gemma":0.0006760867,"teacher_disagreement_score":0.96910405,"about_ca_system_score_codex":0.0011934493,"about_ca_system_score_gemma":0.0013478087,"threshold_uncertainty_score":0.16339529},"labels":[],"label_agreement":null},{"id":"W1981743768","doi":"10.1109/icpc.2012.6240513","title":"Understanding registration-based abstractions: A quantitative user study","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Programmer; Programming language; Abstraction; Code (set theory); Program comprehension; Comprehension; Software engineering; Software","score_opus":0.23797759878255614,"score_gpt":0.3690136411291174,"score_spread":0.13103604234656124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1981743768","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9978079,0.00002615119,0.0013036149,0.000048891347,0.0000033472115,0.00010349548,0.00008470393,0.00006036475,0.00056153664],"genre_scores_gemma":[0.99596137,0.000050361978,0.0024783658,0.00011058317,0.0000066117327,0.00029371082,0.0001980981,0.00005839919,0.00084250525],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.98961836,0.006315805,0.0007500172,0.001200893,0.0014543455,0.0006605604],"domain_scores_gemma":[0.8654584,0.10500097,0.0045028995,0.0077467947,0.014537241,0.0027536922],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022356097,0.0010427699,0.0010920668,0.0022207587,0.0017415001,0.0032598842,0.0014192746,0.0022800365,0.0029229508],"category_scores_gemma":[0.094806366,0.0008188636,0.00070884073,0.00090080925,0.0018417281,0.0054269484,0.0028236047,0.002099661,0.0011237732],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019656902,0.0036547438,0.24128035,0.00088336767,0.00016379992,0.0011293361,0.6622486,0.0015861847,0.022853697,0.00093351584,0.0026426606,0.06065804],"study_design_scores_gemma":[0.0004587107,0.016724432,0.46940053,0.0005094829,0.00039364872,0.0039046225,0.41717097,0.037910257,0.02534907,0.0026095975,0.024492076,0.0010766214],"about_ca_topic_score_codex":0.0017542845,"about_ca_topic_score_gemma":0.0012783965,"teacher_disagreement_score":0.022356097,"about_ca_system_score_codex":0.0008894272,"about_ca_system_score_gemma":0.00058438693,"threshold_uncertainty_score":0.118231714},"labels":[],"label_agreement":null},{"id":"W1981974420","doi":"10.1006/ijhc.1999.0380","title":"Evaluating a domain-specialist-oriented knowledge management system","year":2000,"lang":"en","type":"article","venue":"International Journal of Human-Computer Studies","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Rotation formalisms in three dimensions; Usability; Human–computer interaction; Domain knowledge; Domain (mathematical analysis); Feature (linguistics); Hierarchy; Knowledge management; Process (computing); Software engineering; Knowledge base; Task (project management); Attractiveness; Artificial intelligence; Engineering; Programming language; Systems engineering; Psychology","score_opus":0.06461842349918082,"score_gpt":0.4004732465168877,"score_spread":0.33585482301770686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1981974420","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98577183,0.00012959626,0.006911511,0.00021046065,0.000024573883,0.0005136852,0.00020462493,0.00042961742,0.0058040875],"genre_scores_gemma":[0.97344285,0.00012471293,0.023322884,0.00013105987,0.000017137116,0.00013548466,0.0007032955,0.000023901694,0.0020984744],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99783164,0.0008085037,0.00027808384,0.00027425075,0.0007021284,0.00010533614],"domain_scores_gemma":[0.98391485,0.009730427,0.0009295292,0.0010069167,0.0032255198,0.0011927272],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038235153,0.0004276775,0.0005225918,0.0014056701,0.00089266023,0.002079458,0.0009947375,0.0011328484,0.0035814866],"category_scores_gemma":[0.021662915,0.00017417359,0.00025637282,0.0009663454,0.0002913881,0.001758849,0.0013491311,0.00041286147,0.00072701945],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.01488807,0.014359019,0.095606364,0.0012679285,0.0007782366,0.00074297644,0.00304839,0.047665156,0.047525063,0.002327559,0.004998971,0.76679224],"study_design_scores_gemma":[0.0055405437,0.03447416,0.2142239,0.00028804303,0.0036975318,0.0010715342,0.0056108143,0.61618227,0.09254335,0.0045089102,0.021540668,0.00031827955],"about_ca_topic_score_codex":0.006661667,"about_ca_topic_score_gemma":0.009182538,"teacher_disagreement_score":0.006661667,"about_ca_system_score_codex":0.0015433127,"about_ca_system_score_gemma":0.002403611,"threshold_uncertainty_score":0.020220935},"labels":[],"label_agreement":null},{"id":"W1982036946","doi":"10.1007/s10664-013-9260-1","title":"An experimental investigation on the effects of context on source code identifiers splitting and expansion","year":2013,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Identifier; Computer science; Source code; Program comprehension; Context (archaeology); Documentation; Acronym; Unique identifier; Internal documentation; Code (set theory); Set (abstract data type); Software documentation; World Wide Web; Information retrieval; Software; Programming language; Software system; Software development; Software development process; Linguistics; Software construction","score_opus":0.021987243302096875,"score_gpt":0.26804226194432085,"score_spread":0.24605501864222398,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982036946","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9935688,0.00009604885,0.002399772,0.00010630603,0.000039476414,0.00021205227,0.00017649347,0.00010338368,0.00329773],"genre_scores_gemma":[0.99205387,0.0000791648,0.0054283463,0.00012401782,0.000036152287,0.00043746462,0.00017511865,0.00011414025,0.0015517663],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.99522734,0.0020110977,0.0005567815,0.0012007197,0.0007560964,0.00024802392],"domain_scores_gemma":[0.78182954,0.18394509,0.012187157,0.015128563,0.0043476634,0.0025619802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004278875,0.000783817,0.00055523013,0.0004662779,0.00093963987,0.0016096131,0.0013180306,0.001092048,0.01003117],"category_scores_gemma":[0.09147315,0.00091106637,0.00029487276,0.0005689462,0.0015758759,0.0026428918,0.0019973086,0.0018062802,0.0007759554],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.07506744,0.082536645,0.1170835,0.0031804345,0.00043632014,0.0010954902,0.028453806,0.007210187,0.4887811,0.01215982,0.0032278828,0.18076737],"study_design_scores_gemma":[0.010762077,0.107458584,0.47916186,0.00055540795,0.0018035788,0.0024371955,0.01644204,0.04659492,0.2815251,0.031836245,0.020695368,0.0007276614],"about_ca_topic_score_codex":0.000703515,"about_ca_topic_score_gemma":0.0009931343,"teacher_disagreement_score":0.01003117,"about_ca_system_score_codex":0.0004268356,"about_ca_system_score_gemma":0.00088787137,"threshold_uncertainty_score":0.033557594},"labels":[],"label_agreement":null},{"id":"W1982206358","doi":"10.1145/2499393.2499397","title":"Using code change types in an analogy-based classifier for short-term defect prediction","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Analogy; Computer science; Classifier (UML); Artificial intelligence; Term (time); Machine learning; Pattern recognition (psychology); Physics; Linguistics","score_opus":0.15347632734310102,"score_gpt":0.34663603760474654,"score_spread":0.19315971026164552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982206358","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6429177,0.00080653816,0.34544957,0.0004651458,0.00021394054,0.0004103459,0.0018676294,0.0048534456,0.003015745],"genre_scores_gemma":[0.88771504,0.00018553532,0.108302906,0.00009638479,0.00012016831,0.00023533872,0.0019841893,0.000079207115,0.0012812307],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99837875,0.00021692138,0.00018350525,0.0004249349,0.00064436626,0.0001515618],"domain_scores_gemma":[0.9902478,0.004910918,0.0012474089,0.0008034484,0.0023595323,0.00043091562],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018560702,0.00096910587,0.0010750301,0.0069561466,0.00059160177,0.0011806784,0.0011718163,0.0018082932,0.0011555695],"category_scores_gemma":[0.011532237,0.00029092317,0.00076519494,0.002919891,0.00049579394,0.002145695,0.00068627443,0.0013492318,0.0011169661],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011289236,0.0011351068,0.280975,0.00029149326,0.00019636002,0.0014076813,0.00043826233,0.08763942,0.024593243,0.0028126133,0.007663901,0.591718],"study_design_scores_gemma":[0.000031192907,0.00036558302,0.03175235,0.000039802275,0.000077639546,0.00069442525,0.0001008657,0.9559397,0.005893053,0.00307709,0.0019697882,0.00005857054],"about_ca_topic_score_codex":0.0040713227,"about_ca_topic_score_gemma":0.003928098,"teacher_disagreement_score":0.0069561466,"about_ca_system_score_codex":0.00055726094,"about_ca_system_score_gemma":0.0007323756,"threshold_uncertainty_score":0.009815991},"labels":[],"label_agreement":null},{"id":"W1982488509","doi":"10.1155/2012/792024","title":"Clustering Methodologies for Software Engineering","year":2012,"lang":"en","type":"article","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":81,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Computer science; Software; Strengths and weaknesses; Task (project management); Software analytics; Software engineering; Data mining; Software system; Software development; Software construction; Data science; Systems engineering; Machine learning; Engineering","score_opus":0.03601315445284729,"score_gpt":0.32493817980490497,"score_spread":0.28892502535205766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982488509","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004171364,0.0037375023,0.9916231,0.0003615257,0.00017744611,0.00018697004,0.00010356689,0.0003777855,0.0030149259],"genre_scores_gemma":[0.016644888,0.004489243,0.9752043,0.00020305217,0.00037360174,0.000639384,0.00038536682,0.00019990148,0.0018602825],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9898732,0.004825537,0.0009053294,0.0013166192,0.0028808198,0.00019861502],"domain_scores_gemma":[0.98977715,0.0062534115,0.0006874963,0.0011974275,0.0019029016,0.00018150386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007809771,0.0026104592,0.0018455083,0.008815504,0.0017546789,0.00365036,0.0032362584,0.0024899004,0.0049728183],"category_scores_gemma":[0.019160314,0.0010472479,0.002172735,0.011421779,0.0029361926,0.0039714095,0.0027151997,0.0034108253,0.003438823],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003717632,0.000079695244,0.0005802613,0.0018316556,0.00027782065,0.0001447922,0.00069652573,0.05340529,0.001494943,0.505104,0.012956632,0.42339122],"study_design_scores_gemma":[0.000037255086,0.00006279836,0.00044442058,0.0004940416,0.00006752642,0.0003441662,0.00027245242,0.13385507,0.0014774656,0.7690418,0.09383785,0.00006509931],"about_ca_topic_score_codex":0.002175703,"about_ca_topic_score_gemma":0.002262708,"teacher_disagreement_score":0.008815504,"about_ca_system_score_codex":0.002724953,"about_ca_system_score_gemma":0.0030823718,"threshold_uncertainty_score":0.041302502},"labels":[],"label_agreement":null},{"id":"W1982621550","doi":"10.1145/2523649.2523650","title":"Uncovering access control weaknesses and flaws with security-discordant software clones","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Computer science; Access control; Software security assurance; Perspective (graphical); Code (set theory); Software; Strengths and weaknesses; Computer security; Information security; Programming language; Artificial intelligence; Biology; Security service; Set (abstract data type); Psychology","score_opus":0.00845662667992501,"score_gpt":0.24002677287450183,"score_spread":0.23157014619457683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982621550","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8872639,0.00029855955,0.11039154,0.00022807873,0.00001246328,0.00013878854,0.00012856803,0.00068617205,0.00085190637],"genre_scores_gemma":[0.9287733,0.000095666495,0.07034339,0.00006607491,0.000010045894,0.00006102264,0.0001879558,0.00009650393,0.00036595346],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9930153,0.0012001437,0.00050938537,0.0018983618,0.0029909816,0.00038579182],"domain_scores_gemma":[0.9075392,0.04301316,0.023762483,0.01614471,0.008481069,0.0010594468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043449206,0.0007671677,0.000993052,0.0050414447,0.0011109281,0.0020937815,0.0015546001,0.0018901889,0.0005970391],"category_scores_gemma":[0.06690099,0.0008659191,0.0009139956,0.0030588924,0.0024204738,0.00464332,0.0023977242,0.0015013324,0.00017680784],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084896764,0.00055605685,0.62859344,0.0005822801,0.00048611118,0.0042218324,0.011336087,0.033287775,0.09094197,0.042307496,0.00096733775,0.18587059],"study_design_scores_gemma":[0.00016048412,0.0010315481,0.18503708,0.0003673378,0.001012415,0.010777979,0.0048123896,0.571534,0.12951455,0.08901923,0.006449479,0.0002834736],"about_ca_topic_score_codex":0.0024132736,"about_ca_topic_score_gemma":0.0028850357,"teacher_disagreement_score":0.0050414447,"about_ca_system_score_codex":0.0012840672,"about_ca_system_score_gemma":0.0016899584,"threshold_uncertainty_score":0.022978425},"labels":[],"label_agreement":null},{"id":"W1982627696","doi":"10.5555/2819009.2819213","title":"Towards generation of software development tasks","year":2015,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Task (project management); Decomposition; Context (archaeology); Process (computing); Software engineering; Task analysis; Software development; Software; Human–computer interaction; Programming language; Systems engineering; Engineering","score_opus":0.10560232098094043,"score_gpt":0.3077589008172072,"score_spread":0.20215657983626678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982627696","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013321907,0.00019099437,0.9783026,0.00045809388,0.00011063988,0.00031952487,0.0007189203,0.0026713705,0.003905958],"genre_scores_gemma":[0.06683587,0.00020828527,0.9261222,0.00010639125,0.000034012144,0.0003069682,0.0029653164,0.0007818982,0.002639119],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995937,0.0015606114,0.00032913857,0.0008653975,0.0010898991,0.00021798782],"domain_scores_gemma":[0.9901743,0.0044002724,0.0006142637,0.0021754692,0.002349758,0.00028597258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029694384,0.001332114,0.00056458375,0.00225983,0.000976762,0.0025258895,0.0016190513,0.0014988469,0.0041570547],"category_scores_gemma":[0.019075822,0.0010619997,0.0022289993,0.0012812964,0.00088963617,0.0023528747,0.0030307842,0.002523472,0.0035517479],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037841543,0.00035950795,0.008434153,0.0014094817,0.00012818015,0.0014669296,0.0052123405,0.07039326,0.026815232,0.13784924,0.025854524,0.72169876],"study_design_scores_gemma":[0.00009863767,0.00018153804,0.0024449334,0.00061746157,0.00010873606,0.00070234446,0.0012124255,0.6730401,0.027482409,0.19143888,0.10257094,0.000101668076],"about_ca_topic_score_codex":0.0026045302,"about_ca_topic_score_gemma":0.0030365775,"teacher_disagreement_score":0.0041570547,"about_ca_system_score_codex":0.0010874945,"about_ca_system_score_gemma":0.003203826,"threshold_uncertainty_score":0.015704095},"labels":[],"label_agreement":null},{"id":"W1982658254","doi":"10.1109/icsme.2014.78","title":"Compiling Clones: What Happens?","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Compiler; Programming language; Source code; Executable; Java; Code refactoring; clone (Java method); Theoretical computer science; Software","score_opus":0.01817148380713891,"score_gpt":0.2600242918082672,"score_spread":0.2418528080011283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982658254","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82869977,0.006738861,0.117006734,0.012961923,0.0009304449,0.00022129552,0.00056369015,0.014439333,0.018437905],"genre_scores_gemma":[0.9561804,0.0018958866,0.0345241,0.0018478718,0.000290449,0.00007336156,0.0003830722,0.0021495218,0.002655347],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9911151,0.0022341744,0.0005357832,0.002293109,0.003010179,0.00081162516],"domain_scores_gemma":[0.95053786,0.026617514,0.0063272016,0.009184032,0.006363486,0.0009698754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045658606,0.0010422778,0.0012354144,0.0022774541,0.0020023496,0.0036452368,0.0011257881,0.0019210131,0.002080904],"category_scores_gemma":[0.06669754,0.00083402114,0.0009674813,0.00286468,0.0034205478,0.0075900527,0.0019676154,0.0017128355,0.0011906622],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006078682,0.00037390608,0.29415408,0.000991513,0.0002488534,0.0064524272,0.015473085,0.010075945,0.029412363,0.017540254,0.01476184,0.6099079],"study_design_scores_gemma":[0.00018276433,0.0014137797,0.34552705,0.002432727,0.0013559825,0.027342234,0.03639277,0.083427206,0.22484775,0.17112945,0.10516063,0.0007877628],"about_ca_topic_score_codex":0.0051387143,"about_ca_topic_score_gemma":0.00353309,"teacher_disagreement_score":0.0051387143,"about_ca_system_score_codex":0.0015426546,"about_ca_system_score_gemma":0.0012173703,"threshold_uncertainty_score":0.024146914},"labels":[],"label_agreement":null},{"id":"W1983378102","doi":"10.1115/detc2009-86270","title":"Performance Comparison of Metamodeling Methods From the Perspective of Sample Quality Merits","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; Western Canada Research Grid","keywords":"Metamodeling; Computer science; Sample (material); Robustness (evolution); Sample size determination; Kriging; Data mining; Machine learning; Artificial intelligence; Mathematics; Statistics","score_opus":0.1250976826116511,"score_gpt":0.4533663654574425,"score_spread":0.3282686828457914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1983378102","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14287417,0.0014592158,0.8521884,0.00029219242,0.000046749665,0.00014850001,0.00017992372,0.0012352522,0.0015756368],"genre_scores_gemma":[0.6411695,0.00064083835,0.35652915,0.000080329286,0.00004077266,0.00022576761,0.0005879368,0.00025551417,0.00047023952],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9915535,0.005401068,0.000560875,0.0006399608,0.0016529072,0.00019168363],"domain_scores_gemma":[0.9125808,0.07132963,0.0031230561,0.0062997346,0.006014724,0.0006520965],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021641916,0.0011065371,0.0016408743,0.0029750483,0.00053095544,0.0015696363,0.001236591,0.0014324806,0.00092406856],"category_scores_gemma":[0.09050243,0.00051309523,0.0013001449,0.0015461611,0.0007743078,0.0025966342,0.0015210644,0.0010256285,0.00023374168],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002276305,0.00022584142,0.021348994,0.00062100973,0.00072294514,0.000104376886,0.0004679275,0.5959185,0.0073118806,0.015648674,0.0010133777,0.35434026],"study_design_scores_gemma":[0.00010602177,0.00026487565,0.005363981,0.00007198395,0.00012515362,0.0001105326,0.00014926249,0.97837424,0.008473512,0.0059838532,0.0009234072,0.000053278363],"about_ca_topic_score_codex":0.0031573735,"about_ca_topic_score_gemma":0.002106144,"teacher_disagreement_score":0.021641916,"about_ca_system_score_codex":0.0010995833,"about_ca_system_score_gemma":0.0012551892,"threshold_uncertainty_score":0.114454746},"labels":[],"label_agreement":null},{"id":"W1983569349","doi":"10.1109/iceccs.2013.41","title":"Runtime Prediction of Failure Modes from System Error Logs","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Failure mode and effects analysis; Computer science; Estimator; Instrumentation (computer programming); Mode (computer interface); Reliability engineering; Data mining; Engineering; Programming language; Statistics; Operating system","score_opus":0.013922945726035052,"score_gpt":0.21350367601257622,"score_spread":0.19958073028654116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1983569349","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56036353,0.00042671265,0.40863365,0.00014771736,0.000042609467,0.000132439,0.0034147578,0.025371836,0.0014667234],"genre_scores_gemma":[0.9416017,0.0001690371,0.054136693,0.000020900035,0.000021493874,0.000065929824,0.0031120838,0.00026707843,0.00060509524],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989091,0.00012294103,0.0000818436,0.00022621914,0.00057528156,0.00008460458],"domain_scores_gemma":[0.9914289,0.0034692604,0.0014550935,0.0016945606,0.0017088524,0.00024325162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001283006,0.0011153027,0.00078817277,0.002231855,0.00026309554,0.00088056177,0.00097441784,0.00048208606,0.00072315265],"category_scores_gemma":[0.010700025,0.0003282987,0.0003779689,0.0011734387,0.00024407123,0.0012505846,0.00050938345,0.00096250226,0.0007264668],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007866848,0.00062433654,0.26583973,0.0003988552,0.00019148564,0.000361447,0.0004533699,0.29844597,0.042619336,0.0011901305,0.004821755,0.38426682],"study_design_scores_gemma":[0.000009877783,0.0001062213,0.032307886,0.000018721763,0.00002050667,0.0001053387,0.000041004987,0.95109355,0.014626575,0.0009554977,0.00068629417,0.000028595863],"about_ca_topic_score_codex":0.005147893,"about_ca_topic_score_gemma":0.007067308,"teacher_disagreement_score":0.005147893,"about_ca_system_score_codex":0.0003839711,"about_ca_system_score_gemma":0.00079171744,"threshold_uncertainty_score":0.010235906},"labels":[],"label_agreement":null},{"id":"W1983705529","doi":"10.1109/wpc.2005.26","title":"On Evaluating the Layout of UML Class Diagrams for Program Comprehension","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Class diagram; Story-driven modeling; Computer science; Unified Modeling Language; Communication diagram; UML tool; Applications of UML; Class (philosophy); Readability; Program comprehension; Programming language; Flowchart; Activity diagram; Software engineering; Perspective (graphical); Software; Software system; Artificial intelligence","score_opus":0.07198969455712151,"score_gpt":0.38810603474630506,"score_spread":0.31611634018918355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1983705529","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6635483,0.0021585154,0.3136453,0.0005948742,0.00017445491,0.0017747787,0.00066336757,0.0043129446,0.013127574],"genre_scores_gemma":[0.65427715,0.00090034166,0.33973625,0.00010091524,0.000049118233,0.0007546811,0.001591952,0.0009490523,0.0016404699],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9832977,0.009046803,0.0014252579,0.0009480231,0.0048512747,0.00043090477],"domain_scores_gemma":[0.7082745,0.23784962,0.011401763,0.0057655433,0.034143724,0.0025649038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014726767,0.0016812371,0.0009885136,0.0050054467,0.0011155136,0.0032622819,0.0011613031,0.0015192573,0.004157904],"category_scores_gemma":[0.21164177,0.00036171626,0.00061720976,0.003499572,0.00094208715,0.0054426137,0.0013913872,0.0007343076,0.0009553355],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021693835,0.0014486044,0.029053407,0.0026224945,0.00016399067,0.0004263107,0.0073372987,0.026240539,0.056295414,0.006275374,0.0064860834,0.86148113],"study_design_scores_gemma":[0.0014300437,0.015104287,0.20091565,0.0026188383,0.0009871123,0.0014948675,0.014924747,0.46153757,0.23245445,0.024744421,0.04313739,0.00065076415],"about_ca_topic_score_codex":0.0037077072,"about_ca_topic_score_gemma":0.004718425,"teacher_disagreement_score":0.014726767,"about_ca_system_score_codex":0.0016461156,"about_ca_system_score_gemma":0.0018594665,"threshold_uncertainty_score":0.07788354},"labels":[],"label_agreement":null},{"id":"W1983732413","doi":"10.1109/pic.2014.6972376","title":"Requirements verification: Legal challenges in compliance testing","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University; Carleton University","funders":"","keywords":"Audit; Compliance (psychology); Conformance testing; Computer science; Context (archaeology); Task (project management); Ask price; Risk analysis (engineering); Software engineering; Accounting; Business; Engineering; Systems engineering; Psychology; Standardization","score_opus":0.23322567166479255,"score_gpt":0.3380870713217588,"score_spread":0.10486139965696625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1983732413","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02211772,0.011721935,0.56285906,0.36372894,0.0008434032,0.00039372419,0.00019363542,0.00088272954,0.037258796],"genre_scores_gemma":[0.6090806,0.007100664,0.34815642,0.027724244,0.0023630438,0.00067292294,0.00030309416,0.00051009166,0.004088875],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.76952225,0.1436038,0.013854622,0.0099756885,0.059093494,0.0039501903],"domain_scores_gemma":[0.38709474,0.5217876,0.016030954,0.03245306,0.039756782,0.002876893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.15980084,0.00085434964,0.0021619939,0.005159129,0.005063586,0.013712147,0.006371025,0.013274162,0.0031220838],"category_scores_gemma":[0.3983271,0.0014109348,0.0015551435,0.0040974785,0.0316909,0.033201866,0.009392085,0.01469168,0.0010324165],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000066843095,0.000114903894,0.002323469,0.00056561665,0.000041231844,0.00068791804,0.0023061803,0.005372884,0.000677218,0.8791122,0.0060056355,0.10272594],"study_design_scores_gemma":[0.000045529792,0.000069607755,0.0006612194,0.0010917458,0.000022352906,0.0011181246,0.0013601426,0.0128511945,0.000894794,0.957221,0.02459469,0.00006955439],"about_ca_topic_score_codex":0.0058589745,"about_ca_topic_score_gemma":0.0031170112,"teacher_disagreement_score":0.15980084,"about_ca_system_score_codex":0.0052134837,"about_ca_system_score_gemma":0.01340832,"threshold_uncertainty_score":0.84511775},"labels":[],"label_agreement":null},{"id":"W1983733882","doi":"","title":"Why did this code change","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Commit; Automatic summarization; Code (set theory); Documentation; Code review; Source code; Set (abstract data type); Programming language; World Wide Web; Information retrieval; Static program analysis; Database; Software development; Software","score_opus":0.04837730356648695,"score_gpt":0.2792348650036265,"score_spread":0.23085756143713954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1983733882","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49709722,0.005204158,0.06553861,0.22298063,0.019282749,0.0010123764,0.004358001,0.0063411538,0.1781851],"genre_scores_gemma":[0.8824304,0.0012661734,0.02117909,0.016908603,0.0010749344,0.00018756095,0.0018169645,0.001895652,0.07324061],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968791,0.00066770235,0.0002056302,0.0006492672,0.0011969068,0.00040139066],"domain_scores_gemma":[0.9843024,0.0031108172,0.0024223255,0.0018269335,0.0069777714,0.0013598101],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002592123,0.00038417266,0.000290385,0.0011176718,0.0024046951,0.0023279872,0.00062077865,0.0021969092,0.011888649],"category_scores_gemma":[0.024635866,0.00038965474,0.00060456444,0.00090651045,0.0018277734,0.0029123665,0.0015220556,0.0034782933,0.0045094523],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008633277,0.00043675944,0.13819212,0.0011522959,0.00025140567,0.014312341,0.027626203,0.001792511,0.03280676,0.05093945,0.28908864,0.4425382],"study_design_scores_gemma":[0.00009722731,0.00047311862,0.10992426,0.0011155244,0.00027713523,0.0077561154,0.020333696,0.005690927,0.025255196,0.022750057,0.80607444,0.0002523232],"about_ca_topic_score_codex":0.014087634,"about_ca_topic_score_gemma":0.02039612,"teacher_disagreement_score":0.014087634,"about_ca_system_score_codex":0.0030476132,"about_ca_system_score_gemma":0.002988333,"threshold_uncertainty_score":0.039771497},"labels":[],"label_agreement":null},{"id":"W1983866039","doi":"10.5555/2820518.2820590","title":"A dataset of the activity of the git super-repository of Linux in 2012","year":2015,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal; University of Victoria","funders":"","keywords":"Linux kernel; Computer science; World Wide Web; Kernel (algebra); Process (computing); Operating system","score_opus":0.026511254488833996,"score_gpt":0.26521439057043333,"score_spread":0.23870313608159935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1983866039","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049618937,0.000546956,0.0009699727,0.00029250997,0.00009592134,0.00008484401,0.9432911,0.0016572207,0.0034425575],"genre_scores_gemma":[0.017483931,0.000125683,0.0017652198,0.000045566187,0.000021947753,0.00009727862,0.979192,0.00009285516,0.0011755425],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9980057,0.00019600283,0.0003083722,0.0004605524,0.000765857,0.00026356004],"domain_scores_gemma":[0.9934098,0.0010693164,0.0013049253,0.001230112,0.0022089235,0.00077683583],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0009349915,0.00096856867,0.0008114255,0.006771364,0.00077467423,0.0014939258,0.001253228,0.0011698089,0.0034434507],"category_scores_gemma":[0.0066146003,0.0003624724,0.0007937904,0.011183123,0.00035246985,0.001318617,0.0015700946,0.0011273565,0.007058319],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008177292,0.00042828236,0.13062115,0.0018476354,0.00023235509,0.00065870956,0.0014851755,0.002888434,0.0041329674,0.0022037078,0.793766,0.060917865],"study_design_scores_gemma":[0.00013046498,0.00015975912,0.39463818,0.00034038708,0.0000975844,0.0008068338,0.0015384624,0.006532519,0.0037216947,0.0014398821,0.5904621,0.00013222171],"about_ca_topic_score_codex":0.029153619,"about_ca_topic_score_gemma":0.061781835,"teacher_disagreement_score":0.9932286,"about_ca_system_score_codex":0.0011758865,"about_ca_system_score_gemma":0.0015514282,"threshold_uncertainty_score":0.05796784},"labels":[],"label_agreement":null},{"id":"W1983886545","doi":"10.1109/icsm.2011.6080796","title":"An automatic framework for extracting and classifying near-miss clone genealogies","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; clone (Java method); Artificial intelligence; Natural language processing; Gene","score_opus":0.09661343911364732,"score_gpt":0.334641237600227,"score_spread":0.23802779848657968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1983886545","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052263875,0.00053425494,0.93034774,0.00014390718,0.000020678166,0.00024226196,0.0013779582,0.014300888,0.00076836144],"genre_scores_gemma":[0.16833802,0.00019893283,0.8267148,0.000081330334,0.000027710441,0.0002592137,0.002826892,0.00048094048,0.0010722277],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973429,0.00026912516,0.00029595665,0.0009938217,0.0009546595,0.00014354942],"domain_scores_gemma":[0.98858494,0.0042912974,0.0024590215,0.0016833938,0.0027099575,0.00027144206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002310406,0.0007627437,0.0012299323,0.010027203,0.0010850839,0.0017059264,0.0020972327,0.0014400285,0.00087246485],"category_scores_gemma":[0.016486265,0.00067539373,0.0009002953,0.0045325756,0.0010486825,0.003369338,0.0012552674,0.0012181572,0.0006004314],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021076991,0.00024713282,0.04965905,0.0005196045,0.00017124838,0.0005825241,0.0021013862,0.014583297,0.043693405,0.011969124,0.0069765,0.869286],"study_design_scores_gemma":[0.00008574384,0.00037469668,0.060900982,0.00022894527,0.00026419127,0.0034081808,0.0010400044,0.7811946,0.08194498,0.036696788,0.03357889,0.00028192322],"about_ca_topic_score_codex":0.009341988,"about_ca_topic_score_gemma":0.013029894,"teacher_disagreement_score":0.010027203,"about_ca_system_score_codex":0.00089330005,"about_ca_system_score_gemma":0.0019753124,"threshold_uncertainty_score":0.018575191},"labels":[],"label_agreement":null},{"id":"W1984132113","doi":"10.1016/s0950-5849(03)00002-8","title":"An experiment in software component retrieval","year":2003,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Component (thermodynamics); Component-based software engineering; Software; Software engineering; Information retrieval; Data mining; Programming language; Software development; Physics","score_opus":0.01031131112265868,"score_gpt":0.2583414944196067,"score_spread":0.248030183296948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984132113","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97992146,0.00072983943,0.0089400485,0.0012216055,0.0005492754,0.0014938973,0.00082586176,0.0007030025,0.005614998],"genre_scores_gemma":[0.94860417,0.00052989344,0.03313171,0.002157753,0.00046398467,0.0020757103,0.001927444,0.00035694326,0.010752486],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99186337,0.003818529,0.0014545065,0.0012104197,0.0012630032,0.00039015608],"domain_scores_gemma":[0.8052848,0.17497943,0.002653819,0.010967743,0.003631858,0.0024823125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009931232,0.0015322515,0.002099963,0.00062328763,0.0016342836,0.0020018057,0.0027584957,0.00583861,0.012853615],"category_scores_gemma":[0.08069089,0.0015262004,0.00110601,0.00082012627,0.0020477986,0.0077101965,0.0020027985,0.003386019,0.0031762228],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.2366069,0.1459407,0.021082021,0.010089401,0.0012678119,0.0028145078,0.016743867,0.011036144,0.3077119,0.01100801,0.019079281,0.2166194],"study_design_scores_gemma":[0.14223985,0.39009342,0.05859616,0.0009068718,0.0038787075,0.005606998,0.005200404,0.0932637,0.21297668,0.037106503,0.04852977,0.0016008957],"about_ca_topic_score_codex":0.0020395322,"about_ca_topic_score_gemma":0.0013020452,"teacher_disagreement_score":0.012853615,"about_ca_system_score_codex":0.00072318263,"about_ca_system_score_gemma":0.0014845809,"threshold_uncertainty_score":0.052522004},"labels":[],"label_agreement":null},{"id":"W1984271155","doi":"10.1109/msr.2010.5463291","title":"A comparative exploration of FreeBSD bug lifetimes","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Eclipse; Computer science; Software bug; Data mining; Tracking (education); Programming language; Software","score_opus":0.04521716619195558,"score_gpt":0.3119309425609017,"score_spread":0.2667137763689461,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984271155","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9727959,0.003137873,0.012172275,0.0002736489,0.000031861095,0.000036131765,0.0071231453,0.0014322605,0.0029968917],"genre_scores_gemma":[0.9780682,0.0005411549,0.008487766,0.00003046132,0.000018585433,0.000025057921,0.011915927,0.00017906664,0.0007338358],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9976713,0.00058410014,0.00021953674,0.00040002543,0.0009798661,0.0001451348],"domain_scores_gemma":[0.9528081,0.03517913,0.0035929044,0.0027193173,0.0049320837,0.0007684816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004310998,0.00055664405,0.0006599243,0.0152076045,0.00038379015,0.0012151622,0.00073928543,0.00048315476,0.001357088],"category_scores_gemma":[0.028147971,0.00020578361,0.00079466606,0.0064486125,0.0003532722,0.0027162486,0.0009970327,0.0005000367,0.00048864016],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013602683,0.00031632587,0.5590094,0.0010608769,0.0006639321,0.0003641649,0.0018039457,0.0419788,0.004263455,0.002865395,0.00685617,0.37945732],"study_design_scores_gemma":[0.00009005587,0.0016299108,0.6596722,0.00030469056,0.0003250691,0.002082212,0.0020951594,0.2967338,0.009746115,0.0077717397,0.01937514,0.00017383194],"about_ca_topic_score_codex":0.0034412518,"about_ca_topic_score_gemma":0.0059145414,"teacher_disagreement_score":0.0152076045,"about_ca_system_score_codex":0.0005668071,"about_ca_system_score_gemma":0.00043903384,"threshold_uncertainty_score":0.022799015},"labels":[],"label_agreement":null},{"id":"W1984321109","doi":"10.1109/reet.2010.5633110","title":"Planned programming problem gotchas as lessons in requirements engineering","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Experiential learning; Plan (archaeology); Misfortune; Value (mathematics); Artificial intelligence; Mathematics education; Psychology","score_opus":0.02508135679058213,"score_gpt":0.29979718644643916,"score_spread":0.274715829655857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984321109","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17602913,0.0034616024,0.62966603,0.030109692,0.0017299992,0.0008270849,0.00022757464,0.0024942216,0.1554546],"genre_scores_gemma":[0.55169034,0.0025401833,0.40773278,0.0019862107,0.00034042433,0.0004698557,0.00031122082,0.00061757746,0.034311447],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99449277,0.0038344578,0.00015852952,0.00039594396,0.0008818526,0.00023653629],"domain_scores_gemma":[0.9838098,0.013176176,0.0005575952,0.0010767021,0.0006163811,0.0007633582],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003970327,0.0009854509,0.00046057694,0.0007297912,0.0011430286,0.002636986,0.0022994734,0.001629941,0.006562154],"category_scores_gemma":[0.022679446,0.00053715607,0.0007394574,0.000601893,0.004250684,0.00556374,0.0040161503,0.0050431024,0.0014627903],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038504344,0.0011348604,0.002131079,0.0015964867,0.000057560184,0.0019030483,0.0325283,0.021965295,0.006553678,0.3717006,0.060807247,0.4992368],"study_design_scores_gemma":[0.00021659589,0.0010552383,0.0015381912,0.0009417214,0.000033125554,0.0027066756,0.012936005,0.037105307,0.007666442,0.57236487,0.36329696,0.00013890535],"about_ca_topic_score_codex":0.00045097625,"about_ca_topic_score_gemma":0.0014724545,"teacher_disagreement_score":0.006562154,"about_ca_system_score_codex":0.0014955606,"about_ca_system_score_gemma":0.00134778,"threshold_uncertainty_score":0.02195257},"labels":[],"label_agreement":null},{"id":"W1984455889","doi":"10.1002/spip.161","title":"Simulations for very early lifecycle quality evaluations","year":2002,"lang":"en","type":"article","venue":"Software Process Improvement and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"University of British Columbia; West Virginia University; National Aeronautics and Space Administration","keywords":"Computer science; Interdependence; Quality (philosophy); Outcome (game theory); Key (lock); Monte Carlo method; Process (computing); Software engineering; Risk analysis (engineering); Mathematics","score_opus":0.0773490813847675,"score_gpt":0.3976941463140998,"score_spread":0.3203450649293323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984455889","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09441449,0.0004303748,0.89271945,0.0005784548,0.00006262917,0.00021124346,0.00036461363,0.0015889609,0.009629794],"genre_scores_gemma":[0.70977527,0.00022159785,0.28677306,0.00011195867,0.000018846147,0.0005980093,0.00031502207,0.00014456343,0.002041645],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99846715,0.0011290502,0.00005537815,0.00008035633,0.00020021692,0.00006782821],"domain_scores_gemma":[0.97673595,0.021097433,0.0006011363,0.00067998143,0.0007094251,0.00017598087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036363148,0.0007587785,0.0009116704,0.0010152338,0.00046817475,0.0009615702,0.0012438975,0.0014668545,0.008178907],"category_scores_gemma":[0.016601887,0.00062754477,0.00080482254,0.00067367736,0.0007624668,0.0011814,0.0010626349,0.001796186,0.00056515285],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000054168006,0.00003701639,0.0003622084,0.00002755879,0.0000145327,0.00002095317,0.000027320737,0.9857022,0.00022271278,0.008416459,0.00026487277,0.004849975],"study_design_scores_gemma":[0.000015322144,0.000010922759,0.000037415826,0.000006685035,0.0000023148382,0.0000026200153,0.0000036392867,0.9943463,0.0002240941,0.0050351187,0.0003126017,0.0000029746705],"about_ca_topic_score_codex":0.005437723,"about_ca_topic_score_gemma":0.0050366274,"teacher_disagreement_score":0.008178907,"about_ca_system_score_codex":0.0012419181,"about_ca_system_score_gemma":0.001008029,"threshold_uncertainty_score":0.027361214},"labels":[],"label_agreement":null},{"id":"W1984610995","doi":"10.1109/csmr-wcre.2014.6747198","title":"Analysis and clustering of model clones: An automotive industrial experience","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Automotive industry; Cluster analysis; Computer science; Similarity (geometry); Cluster (spacecraft); Data mining; Artificial intelligence; Engineering; Programming language","score_opus":0.049996367135238313,"score_gpt":0.30284790374511833,"score_spread":0.25285153660988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984610995","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8191184,0.00043980032,0.17547986,0.00043870864,0.000016802014,0.0001368035,0.00016412084,0.0020527618,0.0021527107],"genre_scores_gemma":[0.7904008,0.00036642523,0.20570457,0.00010720863,0.000013874667,0.00006147363,0.0006400417,0.0007185176,0.0019871346],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.994507,0.0023229853,0.00026907926,0.00092603423,0.0017559774,0.00021882332],"domain_scores_gemma":[0.9780357,0.01294146,0.0009998191,0.0031486568,0.004450148,0.00042424083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005791962,0.00074343465,0.0008246181,0.002578818,0.0014375314,0.0013254181,0.0017124024,0.001151614,0.0010816451],"category_scores_gemma":[0.023111898,0.00053912314,0.0006277095,0.0024480927,0.001144166,0.002044313,0.0015532958,0.0009337118,0.00029838143],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006213847,0.0013386187,0.098456636,0.00077672204,0.00025047574,0.0017756773,0.04022054,0.06878732,0.08083838,0.00690384,0.005349358,0.694681],"study_design_scores_gemma":[0.00023907502,0.003190724,0.089560784,0.00028387125,0.0004853749,0.0042833434,0.018822053,0.58170754,0.22967063,0.010765628,0.060628228,0.00036285567],"about_ca_topic_score_codex":0.0065927007,"about_ca_topic_score_gemma":0.009614232,"teacher_disagreement_score":0.0065927007,"about_ca_system_score_codex":0.0015633628,"about_ca_system_score_gemma":0.0010352668,"threshold_uncertainty_score":0.030631185},"labels":[],"label_agreement":null},{"id":"W1984696192","doi":"10.1109/compsac.2013.64","title":"Ontology-Based Classification of Non-functional Requirements in Software Specifications: A New Corpus and SVM-Based Classifier","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Support vector machine; Software requirements specification; Classifier (UML); Functional requirement; Non-functional requirement; Requirements analysis; Software engineering; Software requirements; Ontology; Categorization; Functional specification; Artificial intelligence; Software; Data mining; Software development; Software design; Programming language; Software construction","score_opus":0.08735372587399014,"score_gpt":0.2845660035896021,"score_spread":0.19721227771561195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984696192","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61403394,0.003909634,0.31091225,0.0020715035,0.00096324406,0.0012705774,0.03808638,0.009170305,0.01958214],"genre_scores_gemma":[0.48627904,0.0009601306,0.36501798,0.00038576208,0.00025047662,0.0020196964,0.13527754,0.000744707,0.009064686],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99698216,0.0008146063,0.00047705718,0.0006621363,0.00092940946,0.00013467018],"domain_scores_gemma":[0.9906151,0.005397398,0.00036652703,0.0007790861,0.002629753,0.00021206425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025309776,0.0009529945,0.00079542515,0.004196479,0.0010975971,0.0014258166,0.0015775927,0.002007967,0.0034794575],"category_scores_gemma":[0.0095069595,0.00035412898,0.00089399365,0.0032283512,0.00076325407,0.002358792,0.0015208515,0.0015468716,0.0022712152],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058160647,0.0014166693,0.014040954,0.0028362817,0.000221231,0.002436652,0.0030196363,0.01893879,0.062001254,0.007671055,0.10381348,0.78302246],"study_design_scores_gemma":[0.00033144027,0.00046685862,0.05659224,0.00059844676,0.00034664798,0.0026948254,0.005140485,0.7275739,0.044443056,0.0077178045,0.15387765,0.00021669012],"about_ca_topic_score_codex":0.010044984,"about_ca_topic_score_gemma":0.012128405,"teacher_disagreement_score":0.010044984,"about_ca_system_score_codex":0.0013487733,"about_ca_system_score_gemma":0.0016856936,"threshold_uncertainty_score":0.01997304},"labels":[],"label_agreement":null},{"id":"W1984769753","doi":"10.1145/1985441.1985464","title":"Do time of day and developer experience affect commit bugginess?","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":148,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Commit; Computer science; Software; Metadata; Compensating transaction; Operating system; Software engineering; Computer security; Database; Distributed transaction","score_opus":0.03290802081311553,"score_gpt":0.26552674335369436,"score_spread":0.23261872254057883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984769753","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9953067,0.00048452738,0.0009109424,0.0005301181,0.000023969596,0.000016888374,0.00023090247,0.000046364778,0.0024495153],"genre_scores_gemma":[0.9986499,0.00015971501,0.00034963453,0.000055016346,0.000024664718,0.0000112358375,0.00020678299,0.000028825005,0.0005143376],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99588203,0.0017192672,0.00035442293,0.00052405795,0.0011030857,0.00041717832],"domain_scores_gemma":[0.82576865,0.10263802,0.04843125,0.004916314,0.009233774,0.00901201],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037499606,0.0002963216,0.00047036706,0.002001397,0.0005890036,0.0022541229,0.0006457723,0.0010265367,0.0033533354],"category_scores_gemma":[0.067804575,0.00036435435,0.00039196404,0.0024770175,0.0008467355,0.002618864,0.0011861693,0.0010176513,0.00096682797],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012040568,0.00012281127,0.9826178,0.000057090107,0.000105048246,0.000105971056,0.0022370182,0.00019407147,0.00027381338,0.00016589493,0.0007417658,0.013258439],"study_design_scores_gemma":[0.0000044337353,0.00008915899,0.99598086,0.00001599185,0.00002853693,0.00009051546,0.0021568246,0.0005235784,0.00010142008,0.00020441966,0.0007862635,0.000017928925],"about_ca_topic_score_codex":0.004664855,"about_ca_topic_score_gemma":0.007868367,"teacher_disagreement_score":0.004664855,"about_ca_system_score_codex":0.0005684266,"about_ca_system_score_gemma":0.0006233874,"threshold_uncertainty_score":0.019831955},"labels":[],"label_agreement":null},{"id":"W1984893691","doi":"10.1109/icsm.2013.38","title":"Predicting Bugs Using Antipatterns","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":82,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Code refactoring; Computer science; Software bug; Software quality assurance; Product metric; Software metric; Software engineering; Predictive modelling; Software regression; Process (computing); Software quality; Software; Quality assurance; Quality (philosophy); Data mining; Software development; Machine learning; Engineering; Programming language","score_opus":0.02925354915091929,"score_gpt":0.27559218649483697,"score_spread":0.24633863734391767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984893691","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82921,0.00089393946,0.1549839,0.0004607745,0.000090431815,0.00021955985,0.0057847286,0.0050869775,0.003269784],"genre_scores_gemma":[0.9383178,0.00028087516,0.056006562,0.00006148932,0.000043581298,0.00014567887,0.004089444,0.00022293214,0.00083164626],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99477506,0.0010962825,0.0007293617,0.0012038527,0.00193494,0.00026050737],"domain_scores_gemma":[0.920227,0.03961229,0.01999549,0.0068522054,0.011929114,0.001384002],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037674748,0.0012093587,0.00086397864,0.01117695,0.00047750302,0.0015566716,0.0010377349,0.0010566884,0.0013734493],"category_scores_gemma":[0.04187693,0.00046164967,0.000924541,0.005032509,0.00047909672,0.003263473,0.0015639054,0.0008109091,0.00084791513],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001986873,0.00020468568,0.8021374,0.00021664191,0.00021396771,0.00020891053,0.00038156402,0.019447856,0.0028833312,0.00064985664,0.0024964362,0.17096075],"study_design_scores_gemma":[0.000046388435,0.00060498953,0.36817703,0.0001234537,0.0002399925,0.0012601762,0.00043393005,0.6105874,0.007064668,0.006297532,0.005062304,0.00010209491],"about_ca_topic_score_codex":0.0049823737,"about_ca_topic_score_gemma":0.007563298,"teacher_disagreement_score":0.01117695,"about_ca_system_score_codex":0.00056889653,"about_ca_system_score_gemma":0.00070961734,"threshold_uncertainty_score":0.019924581},"labels":[],"label_agreement":null},{"id":"W1984953175","doi":"10.1109/tse.2014.2361131","title":"Replicating and Re-Evaluating the Theory of Relative Defect-Proneness","year":2014,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal; Queen's University","funders":"","keywords":"Computer science; Context (archaeology); Software quality; Replication (statistics); Software; Source code; Software bug; Code review; Quality (philosophy); Software system; Data science; Software development; Software engineering; Statistics; Programming language","score_opus":0.02886620484135756,"score_gpt":0.2784184521015144,"score_spread":0.24955224726015685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1984953175","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21048322,0.032325767,0.6631627,0.042817447,0.010938121,0.0020247009,0.0012418733,0.0013859224,0.035620272],"genre_scores_gemma":[0.81733394,0.0070157624,0.15777966,0.008819358,0.0030212814,0.0016147621,0.0006239606,0.0007286209,0.0030624978],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.82977146,0.11724895,0.0062173503,0.018078059,0.027400767,0.0012833439],"domain_scores_gemma":[0.25004303,0.5635868,0.018999297,0.12691775,0.038364023,0.0020890916],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.2518065,0.0030797662,0.0031177802,0.008656621,0.0018425796,0.007104244,0.011806192,0.0041805306,0.0047626756],"category_scores_gemma":[0.62056375,0.00091793336,0.0045053894,0.0054335953,0.015778169,0.015913812,0.005980115,0.009646421,0.0013845625],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012902151,0.0007503312,0.074156605,0.0042402917,0.006902615,0.00079141784,0.014812946,0.04776968,0.0026729032,0.44136235,0.014739631,0.39051098],"study_design_scores_gemma":[0.0006068012,0.0023936925,0.030827474,0.0030912142,0.0019582906,0.00058558903,0.004872724,0.079003386,0.003922858,0.8239341,0.048278697,0.0005251677],"about_ca_topic_score_codex":0.009191635,"about_ca_topic_score_gemma":0.0036856756,"teacher_disagreement_score":0.7481935,"about_ca_system_score_codex":0.006406179,"about_ca_system_score_gemma":0.005826174,"threshold_uncertainty_score":0.9226558},"labels":[],"label_agreement":null},{"id":"W1985364899","doi":"10.1007/11880240_18","title":"Semantic Variations Among UML StateMachines","year":2006,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Semantics (computer science); Programming language; Unified Modeling Language; Operational semantics; Formal semantics (linguistics); Software","score_opus":0.012046240363190424,"score_gpt":0.25105087878341886,"score_spread":0.23900463842022843,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985364899","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02706055,0.00066210085,0.9477716,0.0005021161,0.00024349941,0.00011018115,0.0004823071,0.0044800066,0.01868755],"genre_scores_gemma":[0.48315513,0.0009836574,0.49850738,0.00030558713,0.00020606449,0.00024203553,0.0016542675,0.0031301896,0.011815677],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99551255,0.00148204,0.0004869352,0.000659969,0.0016127961,0.0002456889],"domain_scores_gemma":[0.9957016,0.002211382,0.00029684906,0.0011585957,0.0005447137,0.00008684215],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030992932,0.0008558097,0.00066276564,0.002591063,0.0012310429,0.004045827,0.0016393516,0.0014870862,0.003716027],"category_scores_gemma":[0.008566054,0.0013977116,0.0013586929,0.0030458998,0.0024935424,0.008358944,0.0028295885,0.0025666605,0.0011036588],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000068118,0.000031334625,0.00042437192,0.00012994904,0.00002210418,0.00025669349,0.0015788643,0.0044611515,0.0045017973,0.89939976,0.002375287,0.08675061],"study_design_scores_gemma":[0.000020995425,0.00003956636,0.0004862477,0.00019496949,0.00010512911,0.0004648274,0.00030160756,0.046751264,0.016513946,0.8401945,0.09486789,0.000059076003],"about_ca_topic_score_codex":0.0014906838,"about_ca_topic_score_gemma":0.0021042798,"teacher_disagreement_score":0.004045827,"about_ca_system_score_codex":0.001678733,"about_ca_system_score_gemma":0.0009766606,"threshold_uncertainty_score":0.0163908},"labels":[],"label_agreement":null},{"id":"W1985376335","doi":"10.1109/csmr.2012.25","title":"A Market-Based Bug Allocation Mechanism Using Predictive Bug Lifetimes","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software bug; Task (project management); Eclipse; Software regression; Software; Schedule; Process (computing); Software engineering; Software development; Software quality; Operating system; Engineering; Systems engineering","score_opus":0.02033430767841457,"score_gpt":0.2650024351437539,"score_spread":0.2446681274653393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985376335","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06931537,0.00032550428,0.9231501,0.0007247594,0.00014819634,0.00053556246,0.000099862766,0.003028154,0.00267246],"genre_scores_gemma":[0.83756423,0.00009698176,0.15935336,0.00015046177,0.000101063044,0.00037175615,0.00010906968,0.000120502984,0.0021325604],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941713,0.0024086745,0.00038490453,0.0012709465,0.001262627,0.00050140027],"domain_scores_gemma":[0.97678727,0.010911027,0.003963036,0.003164653,0.003510878,0.0016631982],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01881506,0.0010160892,0.0019736025,0.0030486218,0.0014901445,0.002804484,0.0064515984,0.0019443404,0.004360999],"category_scores_gemma":[0.029585492,0.0010164455,0.0011505872,0.0018650354,0.001448054,0.0049570515,0.0023851423,0.0016485378,0.00091555575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013797814,0.0017693712,0.01701815,0.00042509355,0.0005495823,0.0004929122,0.0007245962,0.40940148,0.015431339,0.075028226,0.011548176,0.46623132],"study_design_scores_gemma":[0.00017789635,0.00028054442,0.0013826138,0.00001291099,0.000062314175,0.00016032421,0.000054636075,0.9795323,0.0019185825,0.014378554,0.0019823045,0.000057025583],"about_ca_topic_score_codex":0.0029413607,"about_ca_topic_score_gemma":0.0025227738,"teacher_disagreement_score":0.01881506,"about_ca_system_score_codex":0.0019834002,"about_ca_system_score_gemma":0.0031548026,"threshold_uncertainty_score":0.09950471},"labels":[],"label_agreement":null},{"id":"W1985500401","doi":"10.1109/scam.2014.13","title":"A Change-Type Based Empirical Study on the Stability of Cloned Code","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Java; clone (Java method); Computer science; Code (set theory); Cloning (programming); Software evolution; Stability (learning theory); Source code; Consistency (knowledge bases); Software system; Programming language; Software; Biology; Genetics; Artificial intelligence; Software construction; Machine learning","score_opus":0.14233801321602263,"score_gpt":0.3565292413072824,"score_spread":0.2141912280912598,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985500401","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9972994,0.00013594994,0.0017545754,0.000030191268,0.000005197507,0.000033001797,0.00021438084,0.000024714871,0.00050251285],"genre_scores_gemma":[0.9982262,0.000047447236,0.0011523496,0.000012868789,0.000007037688,0.00003097868,0.0003363265,0.0000133879585,0.00017343252],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99012417,0.0035346085,0.0011418433,0.0015582594,0.0032853286,0.00035580102],"domain_scores_gemma":[0.74634707,0.176599,0.03999232,0.014752789,0.020116853,0.0021920744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008025148,0.00027741984,0.00034894393,0.0045573637,0.0005885249,0.0010488554,0.00075521524,0.0006755566,0.0009980444],"category_scores_gemma":[0.10414181,0.0002539832,0.000600454,0.003968796,0.0015334067,0.00205057,0.001083852,0.0010466649,0.00028537834],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028338566,0.00016276342,0.9764331,0.0001018665,0.00013353665,0.00021914978,0.0020569141,0.00070368196,0.0020480745,0.00029563744,0.00017957002,0.017382355],"study_design_scores_gemma":[0.000010709526,0.00062297285,0.98782015,0.000028683064,0.00006237921,0.00076104386,0.001578662,0.0058469023,0.0022047276,0.00031038598,0.00072881853,0.000024597084],"about_ca_topic_score_codex":0.0014302104,"about_ca_topic_score_gemma":0.00164474,"teacher_disagreement_score":0.008025148,"about_ca_system_score_codex":0.0006472092,"about_ca_system_score_gemma":0.00040575478,"threshold_uncertainty_score":0.042441607},"labels":[],"label_agreement":null},{"id":"W1985812401","doi":"10.1145/1146269.1146276","title":"Software diversity","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Diversity (politics); Software engineering; Sociology; Anthropology","score_opus":0.019288555549738094,"score_gpt":0.24336064741036115,"score_spread":0.22407209186062305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985812401","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051909193,0.011994075,0.6953493,0.010067894,0.0014348301,0.00032403835,0.0005194306,0.0031529753,0.22524832],"genre_scores_gemma":[0.7089547,0.00898102,0.18855311,0.003426032,0.0015058182,0.00057984964,0.0015682316,0.0011953197,0.08523594],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9942115,0.0013357368,0.000321954,0.0013318851,0.0023426327,0.00045634992],"domain_scores_gemma":[0.98363274,0.004932327,0.0013382768,0.0051236167,0.0036004102,0.0013726352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004674174,0.0008480096,0.0008403535,0.0021468296,0.0019673298,0.0037201035,0.0020514415,0.001395923,0.012100683],"category_scores_gemma":[0.018431248,0.0005722206,0.0010876023,0.0016500611,0.0035547887,0.005255874,0.007387896,0.0023163909,0.005834933],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018456337,0.00009889936,0.0037638792,0.00044428074,0.000116135874,0.00028754762,0.0014400379,0.011278012,0.006930595,0.64271957,0.017641455,0.31509498],"study_design_scores_gemma":[0.00005847582,0.00028732483,0.00105219,0.00022182148,0.0000891904,0.0012113871,0.00037928578,0.016888909,0.005847864,0.48209816,0.49178842,0.000076932265],"about_ca_topic_score_codex":0.00070460944,"about_ca_topic_score_gemma":0.0006249731,"teacher_disagreement_score":0.012100683,"about_ca_system_score_codex":0.0017577644,"about_ca_system_score_gemma":0.002479698,"threshold_uncertainty_score":0.040480852},"labels":[],"label_agreement":null},{"id":"W1985891579","doi":"10.1007/s10270-013-0325-9","title":"VPML: an approach to detect design patterns of MOF-based modeling languages","year":2013,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; IBM (Canada)","funders":"","keywords":"Computer science; Unified Modeling Language; Model transformation; Software design pattern; Eclipse; Modeling language; Programming language; Semantics (computer science); Model-driven architecture; Structural pattern; Software engineering; Artificial intelligence; Software development; Software design; Consistency (knowledge bases); Software","score_opus":0.05169012021342602,"score_gpt":0.2724063294857177,"score_spread":0.22071620927229169,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985891579","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008104509,0.0000634914,0.9633688,0.00015190455,0.000034959943,0.00025277125,0.0010835459,0.026005683,0.00093436224],"genre_scores_gemma":[0.060023904,0.00011843465,0.93213177,0.00014054164,0.00002175443,0.00039236105,0.0028947173,0.0030914005,0.0011851549],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995605,0.00096582633,0.00050059316,0.00063460093,0.0019742297,0.00031976786],"domain_scores_gemma":[0.9892679,0.005451345,0.0014808247,0.0018924428,0.0016305718,0.00027682004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036059823,0.0020634003,0.0010056328,0.005756905,0.0013104253,0.003972079,0.003345561,0.00275781,0.0044431794],"category_scores_gemma":[0.018121073,0.001962061,0.002726157,0.0022548286,0.0012115043,0.005877345,0.0036070647,0.0034088064,0.0017675877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008357763,0.001025502,0.034723166,0.002871594,0.00073239225,0.0021673944,0.0059717395,0.048839036,0.052405085,0.118266426,0.037847087,0.6943148],"study_design_scores_gemma":[0.000192879,0.00026627583,0.004036218,0.000675438,0.00046568652,0.0015137816,0.0013163715,0.7156453,0.078478895,0.07821576,0.11896147,0.00023203893],"about_ca_topic_score_codex":0.0063733556,"about_ca_topic_score_gemma":0.0071659875,"teacher_disagreement_score":0.0063733556,"about_ca_system_score_codex":0.0011170673,"about_ca_system_score_gemma":0.0028507796,"threshold_uncertainty_score":0.019070446},"labels":[],"label_agreement":null},{"id":"W1985892950","doi":"10.1007/s10664-015-9371-y","title":"An empirical study of integration activities in distributions of open source software","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; Queen's University; Polytechnique Montréal","funders":"","keywords":"Computer science; Component (thermodynamics); Reuse; System integration; Software engineering; Software; Context (archaeology); Software quality; Component-based software engineering; Software development; Systems engineering; Engineering; Operating system","score_opus":0.06030563834253202,"score_gpt":0.3592406911326763,"score_spread":0.2989350527901443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985892950","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9986883,0.000030548123,0.0005973058,0.000036575755,0.0000011094274,0.000007749917,0.000023387818,0.000008472551,0.00060661946],"genre_scores_gemma":[0.999102,0.000023767665,0.00053677184,0.0000062789854,0.000004615678,0.000006629161,0.00006471445,0.000006426392,0.0002487975],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9966625,0.0017515981,0.00017971275,0.00031392806,0.0007977767,0.00029436577],"domain_scores_gemma":[0.8242804,0.1320251,0.021244299,0.0070063053,0.011331908,0.004112005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004107342,0.0002463054,0.00020382751,0.0026957057,0.00089448906,0.0013958,0.0007967394,0.00068318,0.0023127813],"category_scores_gemma":[0.07991409,0.00025712384,0.0002297931,0.0031578674,0.0012797603,0.0034552964,0.0014599416,0.00128049,0.00043123012],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038512133,0.0014157632,0.9459757,0.000048231595,0.000037274218,0.00027604395,0.0076936064,0.0011501141,0.0019877688,0.00315874,0.00036933448,0.037502356],"study_design_scores_gemma":[0.000038689708,0.00061350607,0.96426946,0.000037921847,0.000040693205,0.00065065554,0.011336148,0.016332874,0.0018870295,0.003173647,0.0015895324,0.000029893781],"about_ca_topic_score_codex":0.0039833426,"about_ca_topic_score_gemma":0.004832301,"teacher_disagreement_score":0.004107342,"about_ca_system_score_codex":0.0011044188,"about_ca_system_score_gemma":0.00087317626,"threshold_uncertainty_score":0.02172196},"labels":[],"label_agreement":null},{"id":"W1985961657","doi":"10.1109/icsm.2010.5609680","title":"Analyzing natural-language artifacts of the software process","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software engineering; Software; Process (computing); Software development process; Natural language; World Wide Web; Software development; Negotiation; Data science; Programming language; Natural language processing","score_opus":0.007074368443343385,"score_gpt":0.2662356869759159,"score_spread":0.2591613185325725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985961657","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4949613,0.0014056035,0.48842993,0.00070803065,0.00007577439,0.0009156482,0.0064615165,0.0018687258,0.005173519],"genre_scores_gemma":[0.6119371,0.0011147427,0.36975196,0.0001731149,0.0001439127,0.0012804837,0.012453092,0.00055582507,0.0025897876],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.99033195,0.004528296,0.0008526132,0.001087886,0.0030060352,0.00019318139],"domain_scores_gemma":[0.875689,0.10073528,0.011548791,0.005183327,0.0063592815,0.00048429408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005430424,0.0007804203,0.00043417542,0.008830319,0.0012802193,0.00248878,0.0011258186,0.0010209652,0.0012730446],"category_scores_gemma":[0.046072856,0.00052932533,0.0006314685,0.0064924727,0.0017217172,0.0037296058,0.0011385455,0.0012533595,0.0005947633],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009917422,0.00142845,0.15333518,0.0064261253,0.00044850996,0.010282621,0.13994211,0.024281276,0.10080396,0.039379228,0.009080547,0.5136003],"study_design_scores_gemma":[0.00030297344,0.00115909,0.32684833,0.002191365,0.00045399187,0.012057521,0.03635286,0.23095748,0.08871864,0.1033983,0.1969671,0.0005923328],"about_ca_topic_score_codex":0.0031875595,"about_ca_topic_score_gemma":0.005335132,"teacher_disagreement_score":0.008830319,"about_ca_system_score_codex":0.0010841747,"about_ca_system_score_gemma":0.0017016465,"threshold_uncertainty_score":0.028719187},"labels":[],"label_agreement":null},{"id":"W1986061503","doi":"10.1002/asi.20084","title":"Dynamic Web log session identification with statistical language models","year":2004,"lang":"en","type":"article","venue":"Journal of the American Society for Information Science and Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Alberta","funders":"","keywords":"Session (web analytics); Computer science; Timeout; Smoothing; Identification (biology); Language model; Key (lock); Statistical model; Machine learning; Data mining; Artificial intelligence; World Wide Web","score_opus":0.0074553583984219825,"score_gpt":0.28364079894375055,"score_spread":0.27618544054532856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1986061503","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02633392,0.00009638051,0.968365,0.0001574776,0.000029648256,0.00006932931,0.00015590082,0.0042616692,0.00053067657],"genre_scores_gemma":[0.5642612,0.00012888908,0.43203,0.00014645873,0.00010863055,0.0002468058,0.00068562763,0.00037840736,0.0020139986],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99482393,0.0022455193,0.00027352397,0.00087553513,0.0015179155,0.000263703],"domain_scores_gemma":[0.9781055,0.0137368115,0.0021513507,0.0031367873,0.0024774636,0.0003920747],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005447479,0.0008604501,0.00106583,0.0030449051,0.0005735898,0.0020973703,0.0024618802,0.0012140614,0.001306052],"category_scores_gemma":[0.022226738,0.0005603114,0.0010497526,0.0017916553,0.00073529035,0.004419526,0.001369008,0.0019302865,0.0010650032],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010232523,0.00090936106,0.016346937,0.00028063197,0.0003752138,0.00034284205,0.000753475,0.3018296,0.023934575,0.02859791,0.007118508,0.6184877],"study_design_scores_gemma":[0.0000089052255,0.00003769313,0.0005301942,0.000006713567,0.000013069977,0.000054559812,0.00003342654,0.9867135,0.0044433363,0.0074194055,0.0007156008,0.00002357422],"about_ca_topic_score_codex":0.003730848,"about_ca_topic_score_gemma":0.0034156085,"teacher_disagreement_score":0.005447479,"about_ca_system_score_codex":0.0010665389,"about_ca_system_score_gemma":0.002155253,"threshold_uncertainty_score":0.028809369},"labels":[],"label_agreement":null},{"id":"W1986077805","doi":"10.1109/esem.2009.5316016","title":"Understanding the use of inheritance with visual patterns","year":2009,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Inheritance (genetic algorithm); Computer science; Program comprehension; Visualization; Set (abstract data type); Comprehension; Class (philosophy); Programmer; Scalability; Multiple inheritance; Programming language; Object (grammar); Metric (unit); Human–computer interaction; Theoretical computer science; Object-oriented programming; Artificial intelligence; Software; Software system; Database","score_opus":0.1873095171239972,"score_gpt":0.31242889293704773,"score_spread":0.12511937581305052,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1986077805","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09383284,0.0005142978,0.89628327,0.0011515111,0.000024204477,0.00008932862,0.00020805745,0.0022144623,0.005682049],"genre_scores_gemma":[0.42438838,0.00072029396,0.5714585,0.00013394195,0.000035169276,0.00012117924,0.0003102888,0.0005972099,0.0022349807],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9992224,0.00033158544,0.00005181732,0.00012571744,0.00021509491,0.00005330069],"domain_scores_gemma":[0.99612635,0.00219392,0.0005131759,0.0005142415,0.00052001147,0.00013226464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012846207,0.0007788468,0.0003388087,0.0022413211,0.00048982725,0.0030977614,0.0008334324,0.001068259,0.002729897],"category_scores_gemma":[0.008701651,0.00050917326,0.00045437415,0.0012615859,0.0010824641,0.006124424,0.0017765175,0.0010282128,0.00042671562],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043884624,0.00020663292,0.021902978,0.0012624551,0.00012119482,0.0018986068,0.03457864,0.031868074,0.08907612,0.19590697,0.010678834,0.6120607],"study_design_scores_gemma":[0.00015860083,0.00032580242,0.023330968,0.0008029278,0.00023199689,0.0049372013,0.009205055,0.34626427,0.057389077,0.43206337,0.12506475,0.00022600713],"about_ca_topic_score_codex":0.0024147744,"about_ca_topic_score_gemma":0.00238329,"teacher_disagreement_score":0.0030977614,"about_ca_system_score_codex":0.00060428446,"about_ca_system_score_gemma":0.00054262235,"threshold_uncertainty_score":0.009132385},"labels":[],"label_agreement":null},{"id":"W1986090683","doi":"10.1109/icsm.2012.6405265","title":"Build system issues in multilanguage software","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Set (abstract data type); Software; Software engineering; Key (lock); Process (computing); Exploratory research; Software system; Operating system; Programming language","score_opus":0.014192246775760604,"score_gpt":0.2826516904983988,"score_spread":0.2684594437226382,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1986090683","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9045747,0.0002955245,0.081000715,0.0028801034,0.000028709168,0.00007380232,0.000018846506,0.00029065274,0.010836949],"genre_scores_gemma":[0.9831654,0.00010310196,0.014730766,0.0002687743,0.000008483868,0.00006271169,0.000018239076,0.0002064931,0.0014360279],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.9653279,0.024277804,0.0014294027,0.0016379606,0.005990362,0.0013365418],"domain_scores_gemma":[0.898898,0.07778735,0.009790109,0.0069700927,0.005303852,0.0012505786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025755161,0.000512818,0.00032322193,0.0016799066,0.005023971,0.004346825,0.0013936795,0.0016415031,0.0018680673],"category_scores_gemma":[0.072938934,0.0009123563,0.00053343805,0.0014894628,0.012288343,0.010382458,0.008307415,0.0030569867,0.0002351735],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013221409,0.00009955478,0.033494297,0.00046268813,0.000046376645,0.0031933412,0.8084091,0.0020038618,0.0094467895,0.093575805,0.0014909726,0.047645],"study_design_scores_gemma":[0.000083079445,0.000703111,0.03417083,0.0011293467,0.00018721851,0.008849899,0.6155869,0.015406646,0.024075018,0.15687326,0.14267726,0.00025744448],"about_ca_topic_score_codex":0.0017569003,"about_ca_topic_score_gemma":0.0032702768,"teacher_disagreement_score":0.025755161,"about_ca_system_score_codex":0.0032432305,"about_ca_system_score_gemma":0.0029061767,"threshold_uncertainty_score":0.136208},"labels":[],"label_agreement":null},{"id":"W1986413471","doi":"10.1145/2807426.2807437","title":"Identifying Test Refactoring Candidates with Assertion Fingerprints","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Boilerplate text; Codebase; Computer science; Assertion; Software engineering; Test (biology); Test case; Programming language; Software maintenance; Software; Software system; Machine learning","score_opus":0.04993975554788687,"score_gpt":0.2957746129745205,"score_spread":0.24583485742663364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1986413471","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7575824,0.0027168733,0.21082583,0.0005448091,0.00032123653,0.0008187844,0.0054968228,0.015133111,0.006560062],"genre_scores_gemma":[0.81785214,0.00054509816,0.16170134,0.00021019817,0.00009917285,0.00028578198,0.013931482,0.0015679646,0.0038068362],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958075,0.00056954485,0.0005380881,0.00085142656,0.0017612986,0.0004720612],"domain_scores_gemma":[0.9641462,0.017801998,0.005838385,0.004147263,0.006605273,0.0014608983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032908109,0.001539243,0.0009485214,0.0077057844,0.00077691465,0.0014471648,0.0018836878,0.0019248136,0.0047263256],"category_scores_gemma":[0.03319247,0.0005085054,0.0013612631,0.003305153,0.0005429529,0.0023268524,0.0013921182,0.00092772435,0.0020871894],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023421347,0.00095785473,0.28699768,0.0020590809,0.00025132197,0.017799575,0.0015362618,0.011924924,0.09145758,0.008089275,0.021864738,0.5547195],"study_design_scores_gemma":[0.0005490837,0.002392362,0.16403742,0.001557454,0.0014642581,0.02344216,0.002473758,0.41699415,0.30288005,0.015381622,0.06837542,0.00045218997],"about_ca_topic_score_codex":0.0022003443,"about_ca_topic_score_gemma":0.0036296418,"teacher_disagreement_score":0.0077057844,"about_ca_system_score_codex":0.00053634134,"about_ca_system_score_gemma":0.0016466139,"threshold_uncertainty_score":0.017403722},"labels":[],"label_agreement":null},{"id":"W1986794977","doi":"10.1002/smr.1655","title":"Guest editorial for the special issue on source code analysis and manipulation, SCAM 2012","year":2014,"lang":"en","type":"editorial","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Beijing; Citation; Library science; Queen (butterfly); Computer science; World Wide Web; China; History","score_opus":0.009387677899137272,"score_gpt":0.28198994553684403,"score_spread":0.2726022676377068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1986794977","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000040892057,0.0031621035,0.0002210022,0.017804813,0.97746277,0.000012557902,0.00005802773,0.00006922792,0.0011686647],"genre_scores_gemma":[0.00036824442,0.0024997753,0.00018195318,0.007884338,0.97842467,0.0000169942,0.000059457285,0.000069728,0.010494726],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949588,0.0005319772,0.00058357575,0.000734173,0.0028341827,0.0003573749],"domain_scores_gemma":[0.9773662,0.006196003,0.0016636048,0.00062370964,0.011142179,0.0030082092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062783905,0.0029513217,0.0032993946,0.0047848155,0.002660524,0.008709515,0.0025718391,0.009283408,0.020946318],"category_scores_gemma":[0.028595658,0.0011211081,0.0023920848,0.0018480417,0.0019477338,0.0038158316,0.0016857041,0.01258758,0.012152366],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030422434,0.000011038416,0.000022362656,0.00009754993,0.000012566345,0.00006693047,0.000004401545,0.000020374833,0.00006831835,0.00015263168,0.9952054,0.004308111],"study_design_scores_gemma":[0.000091362774,0.00004940954,0.0005720445,0.000340673,0.00007516587,0.00029692272,0.000027515558,0.00029967303,0.0002207559,0.0013152885,0.9966898,0.000021169646],"about_ca_topic_score_codex":0.001439982,"about_ca_topic_score_gemma":0.0048767827,"teacher_disagreement_score":0.020946318,"about_ca_system_score_codex":0.0024116323,"about_ca_system_score_gemma":0.0035913612,"threshold_uncertainty_score":0.07007241},"labels":[],"label_agreement":null},{"id":"W1987035533","doi":"10.1145/1391984.1391987","title":"Evaluating the benefits of context-sensitive points-to analysis using a BDD-based implementation","year":2008,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":117,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; University of Waterloo","funders":"","keywords":"Computer science; Heap (data structure); Pointer analysis; Call graph; Java; Pointer (user interface); Programming language; Paddle; Theoretical computer science; Static analysis; Artificial intelligence; Operating system","score_opus":0.22449726015584578,"score_gpt":0.4221343935073883,"score_spread":0.1976371333515425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1987035533","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.256029,0.0010401562,0.71650904,0.00043468407,0.00013841807,0.00029151127,0.0005001547,0.018459378,0.0065975673],"genre_scores_gemma":[0.60130996,0.0003024575,0.39598802,0.0001462303,0.000021241236,0.0001238183,0.0003793913,0.0007469842,0.000981873],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937617,0.0019180984,0.00044434378,0.0005512314,0.002903679,0.0004208953],"domain_scores_gemma":[0.98484576,0.009851998,0.00060625497,0.0026822079,0.001825813,0.00018790684],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044389395,0.0013347734,0.0009597012,0.0018409297,0.0006144313,0.0021188206,0.0023762668,0.0012029146,0.0025557068],"category_scores_gemma":[0.021809414,0.001051126,0.0012164488,0.001522847,0.0013705532,0.003356234,0.0016787898,0.0016183348,0.0005650828],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035745967,0.000685513,0.015997842,0.001441418,0.0004715285,0.00039671746,0.00056577916,0.55773777,0.06500242,0.06154757,0.004244809,0.28833407],"study_design_scores_gemma":[0.00023928605,0.00040106862,0.0011857579,0.00006190225,0.00017264245,0.0001004424,0.000073029085,0.9388827,0.04317279,0.012575103,0.003073983,0.00006131453],"about_ca_topic_score_codex":0.006658626,"about_ca_topic_score_gemma":0.0057504675,"teacher_disagreement_score":0.006658626,"about_ca_system_score_codex":0.001415097,"about_ca_system_score_gemma":0.0025845536,"threshold_uncertainty_score":0.023475647},"labels":[],"label_agreement":null},{"id":"W1987105996","doi":"10.1145/1082983.1083160","title":"A framework for describing and understanding mining tools in software development","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Data science; Premise; User needs; Software; Information needs; Software development; Software engineering; Data mining; World Wide Web","score_opus":0.12289475246367247,"score_gpt":0.29620040797673186,"score_spread":0.17330565551305938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1987105996","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00045337967,0.0012913904,0.99380845,0.0012419808,0.000066835004,0.0003366277,0.00014413508,0.00048909953,0.002167973],"genre_scores_gemma":[0.007332034,0.0010946254,0.9893434,0.0002769402,0.000073804425,0.0008828825,0.00028809122,0.00010094539,0.00060724316],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9722607,0.013715792,0.005450721,0.0024607752,0.004987839,0.0011241839],"domain_scores_gemma":[0.97628164,0.015042635,0.0018960496,0.003042921,0.002908586,0.0008280723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034896575,0.0035431415,0.0028591962,0.015772017,0.0052065565,0.018687608,0.0118044615,0.0094977,0.0035758724],"category_scores_gemma":[0.03501158,0.0026971125,0.005976915,0.017788986,0.01347694,0.026367312,0.0077631474,0.007958538,0.0020942718],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002240371,0.00009018094,0.00061553065,0.000849063,0.00007533764,0.00045231517,0.003770417,0.009849135,0.000758751,0.9241476,0.004432347,0.054936916],"study_design_scores_gemma":[0.000040693438,0.00010368961,0.0004175676,0.0012386811,0.000085664964,0.00072681665,0.0014423011,0.03211614,0.0010193175,0.79886323,0.16381133,0.00013453414],"about_ca_topic_score_codex":0.014531618,"about_ca_topic_score_gemma":0.008243849,"teacher_disagreement_score":0.034896575,"about_ca_system_score_codex":0.0059932694,"about_ca_system_score_gemma":0.0090914145,"threshold_uncertainty_score":0.18455291},"labels":[],"label_agreement":null},{"id":"W1987963388","doi":"10.1007/s10664-015-9370-z","title":"Studying high impact fix-inducing changes","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Change impact analysis; Computer science; Software; Quality (philosophy); Risk analysis (engineering); Data science; Software engineering","score_opus":0.07085915187211694,"score_gpt":0.32367575130540127,"score_spread":0.25281659943328433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1987963388","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9881996,0.00025719873,0.0064390474,0.00010156892,0.000020185344,0.000031275777,0.00026321015,0.0001163987,0.0045715407],"genre_scores_gemma":[0.99571383,0.00012276722,0.0024044267,0.00003108743,0.000011490201,0.000014792356,0.00037883205,0.000036408514,0.0012863771],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99868566,0.00035253353,0.00007328843,0.00026375143,0.00043148806,0.00019324181],"domain_scores_gemma":[0.965533,0.023380185,0.003820332,0.004294248,0.0022710373,0.0007011982],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0015792164,0.0003409708,0.00031384887,0.001168645,0.00048695464,0.00081095315,0.0006311093,0.00070160185,0.0062878877],"category_scores_gemma":[0.026664274,0.00023424345,0.00030807007,0.0014109147,0.00050160347,0.0011456001,0.00055331684,0.0011224283,0.000554653],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002174024,0.0049808677,0.46742672,0.0012868122,0.00063613534,0.0019973668,0.00191822,0.05183289,0.0944741,0.05436312,0.007267329,0.31164238],"study_design_scores_gemma":[0.00028711974,0.003504778,0.72352034,0.00022734515,0.0004977784,0.0026736087,0.0034994094,0.11771883,0.08268793,0.04907685,0.016188268,0.000117732154],"about_ca_topic_score_codex":0.0010804763,"about_ca_topic_score_gemma":0.0020876564,"teacher_disagreement_score":0.9984208,"about_ca_system_score_codex":0.00056182913,"about_ca_system_score_gemma":0.00048284561,"threshold_uncertainty_score":0.021035135},"labels":[],"label_agreement":null},{"id":"W1988252511","doi":"10.1109/saner.2015.7081873","title":"TextRank based search term identification for software change tasks","year":2015,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Task (project management); Software; Term (time); Software maintenance; Identification (biology); Domain (mathematical analysis); Software evolution; Software development; Recall; Precision and recall; Software system; Code (set theory); Program comprehension; Artificial intelligence; Information retrieval; Software construction; Programming language; Engineering","score_opus":0.12745973206216427,"score_gpt":0.35199616360569225,"score_spread":0.22453643154352798,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988252511","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31405485,0.011593779,0.6144104,0.0013257923,0.00070279936,0.001315594,0.009144609,0.040597815,0.006854422],"genre_scores_gemma":[0.4575842,0.0014130902,0.5105129,0.0002457734,0.0005247191,0.0006220384,0.018826975,0.00078676787,0.009483557],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970806,0.00091678096,0.00035601892,0.0004931834,0.00097184954,0.00018161985],"domain_scores_gemma":[0.98989,0.005821449,0.0010444811,0.000841225,0.0020273996,0.0003754986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019937342,0.0015557434,0.0015528254,0.009844114,0.0010107461,0.0011443949,0.0014434768,0.0016609684,0.0039062724],"category_scores_gemma":[0.013535905,0.00033081818,0.0009711574,0.0050980765,0.00042210313,0.0031770077,0.000947871,0.0011474774,0.00463917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018042992,0.0008365846,0.01141878,0.0023695838,0.00017383478,0.00047623762,0.000696977,0.011869254,0.069041446,0.0023781562,0.041771162,0.85716367],"study_design_scores_gemma":[0.0005917436,0.0015816465,0.02614695,0.0001943047,0.00041814565,0.0024078274,0.00091376796,0.8492286,0.06996151,0.01256476,0.035735104,0.00025573847],"about_ca_topic_score_codex":0.0037164083,"about_ca_topic_score_gemma":0.0091805775,"teacher_disagreement_score":0.009844114,"about_ca_system_score_codex":0.0008450545,"about_ca_system_score_gemma":0.0017589638,"threshold_uncertainty_score":0.013067782},"labels":[],"label_agreement":null},{"id":"W1988691616","doi":"10.1145/2597073.2597120","title":"Co-evolution of project documentation and popularity within github","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Documentation; Popularity; Reputation; Computer science; Reputation management; World Wide Web; Software documentation; Relation (database); Software; Software development; Database; Software development process; Political science","score_opus":0.01649266624549378,"score_gpt":0.29982254498784433,"score_spread":0.28332987874235055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988691616","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99367094,0.00065307936,0.0027062779,0.00023382228,0.000021390715,0.000026804391,0.00026636093,0.00031145412,0.002109813],"genre_scores_gemma":[0.9954041,0.00019938641,0.0022759712,0.000025301797,0.00003206039,0.000029637953,0.00056810223,0.00019561496,0.0012697804],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98860615,0.003931882,0.000972024,0.0019929397,0.003653103,0.00084388076],"domain_scores_gemma":[0.7960295,0.091295905,0.055035412,0.01519917,0.032896254,0.009543724],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.011857087,0.0004322821,0.0005013597,0.00952653,0.0009995855,0.0037009863,0.0011703745,0.0008347442,0.0013316238],"category_scores_gemma":[0.11242279,0.00049735093,0.00033167808,0.00908036,0.0012636026,0.0043189838,0.003257234,0.0015948861,0.00070354086],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022571818,0.00011613115,0.91627055,0.00021590537,0.000104360915,0.0003618501,0.007342102,0.0009293723,0.002198921,0.0009555791,0.0014410195,0.06983853],"study_design_scores_gemma":[0.0000136093995,0.00019195104,0.9850246,0.00007726186,0.000048970833,0.00059478777,0.0028503821,0.0044063246,0.001045478,0.000533472,0.005150818,0.000062489315],"about_ca_topic_score_codex":0.004084765,"about_ca_topic_score_gemma":0.008118029,"teacher_disagreement_score":0.99047345,"about_ca_system_score_codex":0.0011447885,"about_ca_system_score_gemma":0.0012585747,"threshold_uncertainty_score":0.06270707},"labels":[],"label_agreement":null},{"id":"W1988757916","doi":"10.1109/icsm.2010.5609581","title":"Recovering traceability links between unit tests and classes under test: An improved method","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Universiteit Antwerpen; Technische Universiteit Delft; University of Calgary","keywords":"Traceability; Unit testing; Computer science; Regression testing; Documentation; Benchmark (surveying); Test case; Consistency (knowledge bases); Software engineering; Task (project management); Test suite; Identification (biology); Source code; Class (philosophy); Data mining; Software; Programming language; Software system; Machine learning; Artificial intelligence; Regression analysis; Engineering; Systems engineering; Software construction","score_opus":0.040246712709033954,"score_gpt":0.34526322737796783,"score_spread":0.30501651466893387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988757916","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022956408,0.000106050444,0.9779382,0.00010338702,0.000042236486,0.00017094704,0.00021479733,0.018816605,0.00031220476],"genre_scores_gemma":[0.0582988,0.00017528453,0.93242466,0.00017561432,0.00009148328,0.00033947628,0.0015885689,0.003222359,0.003683712],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9924683,0.0011411277,0.00081911584,0.0018520636,0.0033405425,0.00037890938],"domain_scores_gemma":[0.98512334,0.0060091903,0.0014696746,0.0038115226,0.0032742312,0.0003119827],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003209296,0.0024085853,0.0016024177,0.008328592,0.0008873555,0.0030471715,0.0031561793,0.002373979,0.007164043],"category_scores_gemma":[0.0200313,0.0009535417,0.002927794,0.0029675923,0.0012703569,0.0038256464,0.0029093828,0.0031912285,0.0036235014],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002368437,0.0003890511,0.0076414137,0.0005182445,0.00018868425,0.00046118643,0.0005703349,0.015421366,0.02665909,0.0069252555,0.005826463,0.935162],"study_design_scores_gemma":[0.00024561337,0.00039019733,0.005031788,0.0003024607,0.00041263213,0.0021460825,0.00032151426,0.8517543,0.069864005,0.022164203,0.047064718,0.00030239383],"about_ca_topic_score_codex":0.0091552595,"about_ca_topic_score_gemma":0.00519845,"teacher_disagreement_score":0.0091552595,"about_ca_system_score_codex":0.0011069246,"about_ca_system_score_gemma":0.004381461,"threshold_uncertainty_score":0.023966074},"labels":[],"label_agreement":null},{"id":"W1988795613","doi":"10.1109/itcc.2005.193","title":"Metrics suite for class complexity","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Cyclomatic complexity; Suite; Computer science; Software metric; Software engineering; Class (philosophy); Test suite; Software; Focus (optics); Programming complexity; Software development; Software construction; Test case; Machine learning; Programming language; Artificial intelligence; Regression analysis","score_opus":0.08144730066026003,"score_gpt":0.3238077996962984,"score_spread":0.2423604990360384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988795613","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016043196,0.001895884,0.93852586,0.0009139271,0.00032538973,0.0017351991,0.013303416,0.0162773,0.010979851],"genre_scores_gemma":[0.089907795,0.0007927937,0.87908196,0.00017759686,0.00012800047,0.0034558016,0.02131233,0.0018332204,0.003310553],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9731501,0.0076215765,0.0051413216,0.0015379271,0.012003964,0.0005451476],"domain_scores_gemma":[0.93426037,0.028767789,0.007116966,0.0074806274,0.021148764,0.0012255937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011153113,0.0028643054,0.0021688908,0.016750263,0.0011068307,0.0031264368,0.00248036,0.001382459,0.006669175],"category_scores_gemma":[0.095673956,0.00069056876,0.0018869567,0.012890969,0.0008394139,0.004408604,0.0022839622,0.0022546623,0.002098939],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039067896,0.00036059917,0.028568827,0.0021256867,0.0007345847,0.00047465728,0.00067960203,0.073269464,0.009621084,0.10867323,0.08428914,0.69081247],"study_design_scores_gemma":[0.00029118772,0.002008996,0.03588576,0.0011190131,0.00056051166,0.0025831272,0.0005014404,0.4251577,0.028536439,0.18222545,0.32054266,0.00058769586],"about_ca_topic_score_codex":0.0042990553,"about_ca_topic_score_gemma":0.003975533,"teacher_disagreement_score":0.016750263,"about_ca_system_score_codex":0.0020098495,"about_ca_system_score_gemma":0.0042984327,"threshold_uncertainty_score":0.05898404},"labels":[],"label_agreement":null},{"id":"W1989015001","doi":"10.1109/iwsm.mensura.2014.48","title":"Software Estimation: Transforming Dust into Pots of Gold?","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Agile software development; Craft; Dark Ages; Alchemy; Software; Extreme programming practices; Estimation; Computer science; COCOMO; Software engineering; Engineering; Art; Software development; Art history; Visual arts; Software development process; Systems engineering; Astronomy; Programming language; Physics; Software construction","score_opus":0.011712845384374095,"score_gpt":0.2560155053073448,"score_spread":0.2443026599229707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989015001","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15922095,0.031059625,0.4361856,0.24285002,0.0024974872,0.000126416,0.0002767287,0.0027035943,0.12507954],"genre_scores_gemma":[0.8711922,0.009510754,0.079579584,0.011922692,0.0009231291,0.00007249943,0.0001467557,0.0011832656,0.025468998],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98514026,0.006801698,0.0006901147,0.0014698415,0.0053343903,0.00056373293],"domain_scores_gemma":[0.96998245,0.016126405,0.0018879632,0.005593167,0.0052620573,0.0011481125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016005512,0.0007423359,0.000771126,0.0022445675,0.0028768417,0.009497539,0.001720064,0.0021335192,0.0040716785],"category_scores_gemma":[0.07035473,0.00093714986,0.00045100035,0.0018614013,0.020489492,0.01753796,0.0076686284,0.0055641043,0.002145838],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020120322,0.00011124029,0.012964916,0.0006709275,0.0001004087,0.00028023846,0.039314717,0.0030244202,0.0026135738,0.42856735,0.04055999,0.47159106],"study_design_scores_gemma":[0.000055613422,0.00023251706,0.006673967,0.0014128053,0.00005941663,0.00060733524,0.026168482,0.00906005,0.0044105984,0.5886797,0.36250213,0.00013739469],"about_ca_topic_score_codex":0.0082461685,"about_ca_topic_score_gemma":0.0073800636,"teacher_disagreement_score":0.016005512,"about_ca_system_score_codex":0.0030094148,"about_ca_system_score_gemma":0.00372998,"threshold_uncertainty_score":0.084646285},"labels":[],"label_agreement":null},{"id":"W1989049987","doi":"10.1109/saner.2015.7081841","title":"Cross-project build co-change prediction","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Queen's University","funders":"","keywords":"Computer science; Executable; Eclipse; Construct (python library); Source code; Software; Data mining; Predictive modelling; Machine learning; Code (set theory); Software engineering; Artificial intelligence; Programming language","score_opus":0.10781120010546208,"score_gpt":0.3675787002011224,"score_spread":0.2597675000956603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989049987","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81405056,0.0060855174,0.12637533,0.0012932175,0.0007529608,0.00056171376,0.020628717,0.022097,0.00815493],"genre_scores_gemma":[0.859131,0.0007827031,0.06923043,0.0002745532,0.00014761623,0.00039359898,0.06308365,0.0010320779,0.0059244027],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9930715,0.0010887467,0.00042229163,0.0029981965,0.0018818144,0.0005374645],"domain_scores_gemma":[0.9802161,0.007677275,0.0020182924,0.0041265776,0.0051238067,0.0008380286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068560527,0.0028402738,0.0015314529,0.0054305964,0.00089476095,0.0018782994,0.0023603917,0.0017937679,0.0013303346],"category_scores_gemma":[0.022624776,0.0007360304,0.001688044,0.003814764,0.0006914816,0.0040078675,0.0028750177,0.0033553059,0.0021088559],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009001919,0.00088906015,0.40099183,0.00060398167,0.0007490554,0.00074566953,0.0005952217,0.10065397,0.006491994,0.00088982057,0.054131817,0.4323574],"study_design_scores_gemma":[0.00008320099,0.0004422427,0.10292433,0.00011188189,0.00029870588,0.00071083376,0.00039861785,0.8523436,0.015009524,0.0026861376,0.024884814,0.00010606614],"about_ca_topic_score_codex":0.008272564,"about_ca_topic_score_gemma":0.015555334,"teacher_disagreement_score":0.008272564,"about_ca_system_score_codex":0.0011694067,"about_ca_system_score_gemma":0.0011435639,"threshold_uncertainty_score":0.036258698},"labels":[],"label_agreement":null},{"id":"W1989130222","doi":"10.1109/wcre.2007.15","title":"Clone Detection via Structural Abstraction","year":2007,"lang":"en","type":"article","venue":"Proceedings - Working Conference on Reverse Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Abstraction; Programming language; Computer science; Compiler; Cloning (programming); clone (Java method); Syntax; Representation (politics); Intermediate language; Code (set theory); Theoretical computer science; Abstract syntax tree; Code generation; Line (geometry); Algorithm; Artificial intelligence; Parsing; Mathematics; Operating system; Biology; Genetics","score_opus":0.022816790624592048,"score_gpt":0.25382463284043166,"score_spread":0.23100784221583961,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989130222","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034160722,0.00021969265,0.9575054,0.00010643194,0.000037961745,0.00009841045,0.00016232881,0.0061463583,0.0015626976],"genre_scores_gemma":[0.23636001,0.00022818182,0.75778145,0.00016536066,0.000048625585,0.00019396369,0.0010853055,0.0012010476,0.0029361194],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953637,0.0008703571,0.00029148,0.0010003209,0.002149606,0.00032447404],"domain_scores_gemma":[0.98396033,0.006027208,0.002437747,0.004300371,0.002989181,0.0002850862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024166312,0.0008420484,0.0009097881,0.0048321425,0.0012087069,0.0022000922,0.002154152,0.0012119134,0.0017778304],"category_scores_gemma":[0.019931415,0.00078275776,0.001243227,0.0024919193,0.0017133541,0.0034420092,0.0035715643,0.0015411393,0.0011947346],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036488014,0.00014346998,0.033986535,0.0004633463,0.00017324359,0.0007197058,0.0028951168,0.02129128,0.07697089,0.07043408,0.006851898,0.78570557],"study_design_scores_gemma":[0.00010635482,0.00035968944,0.015516076,0.00019779384,0.00031794078,0.0026610182,0.00067742355,0.5714762,0.18238816,0.16603997,0.060055453,0.00020389896],"about_ca_topic_score_codex":0.0022364906,"about_ca_topic_score_gemma":0.0026395218,"teacher_disagreement_score":0.0048321425,"about_ca_system_score_codex":0.00097280525,"about_ca_system_score_gemma":0.001940855,"threshold_uncertainty_score":0.012780488},"labels":[],"label_agreement":null},{"id":"W1989181588","doi":"10.1145/1852786.1852841","title":"Evaluation of optimized staffing for feature development and bug fixing","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Staffing; Feature (linguistics); Productivity; Computer science; Empirical research; Software bug; Software engineering; Software; Operating system; Mathematics; Statistics; Management","score_opus":0.03652472059376993,"score_gpt":0.3111011085496981,"score_spread":0.27457638795592815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989181588","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98291504,0.00043288732,0.014089599,0.00009609838,0.000023691997,0.00014768785,0.00017141619,0.00039529803,0.0017282754],"genre_scores_gemma":[0.98954356,0.00006590391,0.0096475575,0.00001936341,0.000009718361,0.00008539057,0.00023040002,0.000036701378,0.00036146317],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9912519,0.00483089,0.0006294662,0.00083542365,0.0020732656,0.0003790883],"domain_scores_gemma":[0.88834417,0.08577577,0.011629623,0.006236985,0.0061899945,0.0018234608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011131206,0.000614474,0.00064926984,0.0013682421,0.00032822482,0.0008664061,0.001220529,0.0006771384,0.0015408964],"category_scores_gemma":[0.07959281,0.00036693588,0.00040392258,0.0013035274,0.00081903563,0.0013790414,0.000506044,0.00072443526,0.00024053919],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.018027576,0.004926341,0.115733385,0.00088972226,0.0006693443,0.00014480528,0.0010663137,0.3579701,0.011542509,0.004548918,0.0031015512,0.4813795],"study_design_scores_gemma":[0.0020697664,0.023304883,0.21264078,0.00012075621,0.00048335636,0.00019743321,0.0007178739,0.7407752,0.012764499,0.0035858648,0.003212655,0.00012698301],"about_ca_topic_score_codex":0.0062303613,"about_ca_topic_score_gemma":0.0049773646,"teacher_disagreement_score":0.011131206,"about_ca_system_score_codex":0.0027526673,"about_ca_system_score_gemma":0.002055745,"threshold_uncertainty_score":0.05886817},"labels":[],"label_agreement":null},{"id":"W1989290647","doi":"10.1145/568760.568763","title":"Computational intelligence as an emerging paradigm of software engineering","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software engineering; Software; Data science; Programming language","score_opus":0.03202099377119315,"score_gpt":0.28243543430571333,"score_spread":0.25041444053452017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989290647","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008541796,0.060706504,0.806248,0.025840797,0.00097268756,0.000105901585,0.0002754974,0.00060836703,0.09670049],"genre_scores_gemma":[0.43253496,0.093391426,0.4474091,0.0038487683,0.0035688488,0.00068149646,0.00048149834,0.00022819676,0.017855708],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9977168,0.0009851806,0.00012160537,0.00028361054,0.00077578856,0.00011710276],"domain_scores_gemma":[0.9965359,0.0025096082,0.00015208896,0.00041256793,0.00027081053,0.00011901995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00186873,0.0012030464,0.0013813385,0.0026074601,0.0009129159,0.0064591058,0.0024549644,0.0027366674,0.00270678],"category_scores_gemma":[0.0046747085,0.0005139455,0.00093671866,0.0039781076,0.0066332566,0.007404682,0.002074464,0.0048262943,0.0008160174],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000027628403,0.0000053390613,0.00006766998,0.00009438782,0.0000100323205,0.000018764966,0.0000862073,0.0036506748,0.00006487235,0.9877093,0.0009715609,0.0073185097],"study_design_scores_gemma":[0.0000029196885,0.000007944139,0.000050352835,0.000053485677,0.000004402808,0.000027908114,0.00005476377,0.013666582,0.00006354417,0.96125895,0.024801744,0.000007387421],"about_ca_topic_score_codex":0.0019772698,"about_ca_topic_score_gemma":0.0018607501,"teacher_disagreement_score":0.0064591058,"about_ca_system_score_codex":0.0022781196,"about_ca_system_score_gemma":0.00194518,"threshold_uncertainty_score":0.016528964},"labels":[],"label_agreement":null},{"id":"W1989411104","doi":"10.5555/2486788.2486960","title":"Deciphering the story of software development through frequent pattern mining","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Task (project management); Schema (genetic algorithms); Software; Field (mathematics); World Wide Web; Human–computer interaction; Task analysis; Software development; User interface; Data science; Information retrieval; Programming language; Engineering","score_opus":0.02691618648355044,"score_gpt":0.24873398701257413,"score_spread":0.2218178005290237,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989411104","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5903689,0.0013296949,0.38761312,0.0054226927,0.0001392022,0.0004419991,0.0057085445,0.0020285533,0.0069472287],"genre_scores_gemma":[0.61022913,0.0006283518,0.3831544,0.00022629176,0.00007961224,0.00023523025,0.0037700625,0.00019364816,0.0014833545],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9950488,0.0021459074,0.0005781315,0.0007904767,0.0012481724,0.0001883914],"domain_scores_gemma":[0.9655957,0.022775978,0.004946311,0.0030545115,0.0031064649,0.00052103313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00439672,0.00047893598,0.00041409835,0.0081534255,0.0012218935,0.0029089442,0.0014167501,0.00095156056,0.001062647],"category_scores_gemma":[0.03305277,0.00042911907,0.0005780937,0.008368238,0.0010223537,0.004770312,0.0018104013,0.0012992591,0.0004779413],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048534514,0.000434599,0.22939503,0.0010794606,0.00026968916,0.0029476425,0.01641954,0.012013554,0.013732452,0.012061717,0.016391147,0.6947699],"study_design_scores_gemma":[0.0000807153,0.00042996136,0.22873881,0.0014576713,0.0003579608,0.006484766,0.039917424,0.473969,0.040179424,0.09737276,0.11074429,0.00026723233],"about_ca_topic_score_codex":0.0043613883,"about_ca_topic_score_gemma":0.0064089657,"teacher_disagreement_score":0.0081534255,"about_ca_system_score_codex":0.0009867102,"about_ca_system_score_gemma":0.0010487136,"threshold_uncertainty_score":0.023252368},"labels":[],"label_agreement":null},{"id":"W1989528416","doi":"10.1145/2095654.2095656","title":"Partial models","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Set (abstract data type); Field (mathematics); Class (philosophy); Software; Artificial intelligence; Data mining; Data science; Software engineering; Programming language","score_opus":0.08306788866917031,"score_gpt":0.261379603197753,"score_spread":0.1783117145285827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989528416","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004706986,0.00085684133,0.9535072,0.0014525393,0.00015328382,0.00016682956,0.0015377454,0.0006662925,0.03695226],"genre_scores_gemma":[0.33791012,0.0040758955,0.6087022,0.0010372206,0.00042604524,0.0013848776,0.005380606,0.00053820043,0.040544868],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99684083,0.0010493426,0.00028311842,0.0005756776,0.0010161522,0.00023499383],"domain_scores_gemma":[0.9948861,0.0022108797,0.000419413,0.0014357385,0.00085223553,0.0001956441],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026735768,0.0015699569,0.0011770782,0.0020476202,0.0012561112,0.00508437,0.0030545113,0.0020825209,0.017641654],"category_scores_gemma":[0.00944544,0.0008712598,0.0025204218,0.0018895377,0.0023940532,0.007445525,0.003898219,0.0027442193,0.0037863674],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000024431072,0.00002816304,0.00036141698,0.0001393489,0.000039826817,0.00013535506,0.00019292068,0.038876865,0.00046849585,0.94225097,0.003105488,0.01437669],"study_design_scores_gemma":[0.000024644787,0.000040277024,0.00012560534,0.000086038955,0.000045191704,0.0001678797,0.00008340263,0.13803302,0.000743956,0.8040958,0.05652849,0.0000258091],"about_ca_topic_score_codex":0.0040748008,"about_ca_topic_score_gemma":0.004527546,"teacher_disagreement_score":0.017641654,"about_ca_system_score_codex":0.0016990475,"about_ca_system_score_gemma":0.002119735,"threshold_uncertainty_score":0.05901724},"labels":[],"label_agreement":null},{"id":"W1989571326","doi":"10.1145/1869459.1869518","title":"Refactoring references for library migration","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Declaration; Programming language; Transformation (genetics); Set (abstract data type); Key (lock); Feature (linguistics); Code (set theory); Software engineering; Source code; Software; Computer security; Linguistics","score_opus":0.0269346819908936,"score_gpt":0.2734491500545089,"score_spread":0.24651446806361532,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989571326","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027561065,0.0016525787,0.8999628,0.001196457,0.00084174355,0.00036223934,0.00066774007,0.053343706,0.01441176],"genre_scores_gemma":[0.17878713,0.0011880497,0.78282964,0.00096933293,0.00031052437,0.00030771625,0.0024863516,0.014581,0.018540211],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99387795,0.0019010451,0.00056094176,0.00085020805,0.0022605693,0.0005491743],"domain_scores_gemma":[0.9747756,0.0054187514,0.001554866,0.013136516,0.0045949854,0.0005192249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058768634,0.0013251202,0.0009752208,0.0026041686,0.0015147675,0.002369562,0.0038553781,0.0028069706,0.012849199],"category_scores_gemma":[0.029482722,0.0010755648,0.0015319326,0.0021979797,0.0011661558,0.0060003526,0.0059236265,0.0035810817,0.008087928],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006012883,0.00036635733,0.0066715996,0.0011230079,0.00014442317,0.0018159532,0.0023949747,0.01489446,0.049782477,0.07268839,0.057538263,0.79197866],"study_design_scores_gemma":[0.0002360408,0.00029084116,0.0033854167,0.0010013181,0.00025708583,0.001941416,0.00058948185,0.08858423,0.16343777,0.08843562,0.65149444,0.0003463446],"about_ca_topic_score_codex":0.002091549,"about_ca_topic_score_gemma":0.003043442,"teacher_disagreement_score":0.012849199,"about_ca_system_score_codex":0.0009070146,"about_ca_system_score_gemma":0.0026769966,"threshold_uncertainty_score":0.042984903},"labels":[],"label_agreement":null},{"id":"W1989943253","doi":"10.1145/1321631.1321703","title":"Feature interaction analysis","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Feature (linguistics); Feature model; Software; Software engineering; Code (set theory); Data mining; Programming language; Set (abstract data type)","score_opus":0.013077428793103921,"score_gpt":0.30247123729428743,"score_spread":0.2893938085011835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989943253","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068045944,0.00036011887,0.8779278,0.00025397888,0.00007473028,0.00081359164,0.0036763733,0.008412242,0.040435165],"genre_scores_gemma":[0.5609022,0.00027056425,0.4148112,0.00010530503,0.00004332633,0.0008229122,0.0071301893,0.0008866723,0.015027666],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99749255,0.0003485873,0.00012142994,0.00037790215,0.0013958765,0.00026368946],"domain_scores_gemma":[0.99578756,0.0015928146,0.0003232372,0.0005165015,0.0017040838,0.00007581347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013752058,0.001122669,0.0004979391,0.008033435,0.0010551397,0.0019969826,0.0014522984,0.0008441154,0.014658171],"category_scores_gemma":[0.006189727,0.00030422903,0.0015173531,0.0042932294,0.00062539766,0.0020938634,0.0014071045,0.00070810376,0.0029769766],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005823464,0.00048053593,0.038900305,0.0009170926,0.00029624588,0.0016692104,0.0023432486,0.02775546,0.042856127,0.066417284,0.017912395,0.7998698],"study_design_scores_gemma":[0.00011082287,0.0005770126,0.07517764,0.00032521534,0.0004906144,0.0037422276,0.0033715174,0.4978401,0.102529146,0.113070175,0.20248571,0.00027980583],"about_ca_topic_score_codex":0.0050482466,"about_ca_topic_score_gemma":0.002909694,"teacher_disagreement_score":0.014658171,"about_ca_system_score_codex":0.0009576707,"about_ca_system_score_gemma":0.0012201971,"threshold_uncertainty_score":0.049036503},"labels":[],"label_agreement":null},{"id":"W1990387725","doi":"10.1109/compsacw.2013.93","title":"Analyzing and Predicting Software Quality Trends Using Financial Patterns","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Stock (firearms); Software; Crossover; Stock market; Evolvability; Quality (philosophy); Stock price; Econometrics; Economics; Artificial intelligence; Engineering","score_opus":0.034580476787678165,"score_gpt":0.3040098804700878,"score_spread":0.26942940368240964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1990387725","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.980845,0.00015785325,0.0155592095,0.00015292398,0.000009942496,0.000077980956,0.0012252845,0.00020067344,0.0017712141],"genre_scores_gemma":[0.98393905,0.00016103551,0.014420257,0.000011440191,0.000008362363,0.000037801936,0.0010496153,0.000018817382,0.00035363977],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990928,0.00015592432,0.00015244212,0.0001813981,0.00033408412,0.00008337039],"domain_scores_gemma":[0.9897597,0.0042584497,0.0028164114,0.00060875923,0.002225463,0.00033128468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019012666,0.0004958059,0.0003376218,0.008342971,0.00026456782,0.0013455688,0.0004588226,0.0005293806,0.0007020991],"category_scores_gemma":[0.011916885,0.00021657957,0.00038361072,0.004936397,0.0002420179,0.0019520536,0.00059012714,0.00045391725,0.00034102233],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012412432,0.0001788565,0.8791717,0.0001011108,0.00008676934,0.00019653905,0.00035670924,0.015638066,0.0029212867,0.00071527564,0.0006915223,0.09981799],"study_design_scores_gemma":[0.000013603907,0.00026326292,0.63143533,0.000048819504,0.00008160196,0.00028377632,0.0007887877,0.35888818,0.0041498872,0.0025720946,0.0014330341,0.00004163744],"about_ca_topic_score_codex":0.007439551,"about_ca_topic_score_gemma":0.010934867,"teacher_disagreement_score":0.008342971,"about_ca_system_score_codex":0.0005223216,"about_ca_system_score_gemma":0.00038549432,"threshold_uncertainty_score":0.014792502},"labels":[],"label_agreement":null},{"id":"W1990503178","doi":"10.1142/s0218194002000883","title":"PREDICTING FAULT-PRONE MODULES IN EMBEDDED SYSTEMS USING ANALOGY-BASED CLASSIFICATION MODELS","year":2002,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University","keywords":"Computer science; Software quality; Avionics software; Reliability engineering; Software system; Reliability (semiconductor); Software; Embedded software; Software fault tolerance; Embedded system; Software development; Distributed computing; Fault tolerance; Engineering; Operating system","score_opus":0.03968420817390434,"score_gpt":0.26719730615041987,"score_spread":0.22751309797651553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1990503178","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6589252,0.00023219612,0.33697864,0.00033529577,0.000032418196,0.0002156481,0.00033021116,0.0006803666,0.002269951],"genre_scores_gemma":[0.96359754,0.000056485776,0.03541974,0.000024635234,0.000010823993,0.000057836434,0.00026947726,0.000017717799,0.00054571184],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999183,0.00025993885,0.00007408391,0.0002049623,0.00020125652,0.00007666774],"domain_scores_gemma":[0.9884377,0.009063676,0.0011452365,0.00046562197,0.0006901528,0.00019767288],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023288168,0.0007266174,0.0006453155,0.0027360495,0.00047218555,0.0010313566,0.0012988131,0.0014950032,0.0014254882],"category_scores_gemma":[0.013414819,0.00037333372,0.0007032707,0.001353189,0.0006329239,0.0014616695,0.0006399436,0.0008916521,0.00028339203],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000071048424,0.00016879803,0.021777207,0.000034317676,0.00003381393,0.00017542795,0.00010163163,0.95210147,0.001035944,0.0018136867,0.0003441367,0.022342505],"study_design_scores_gemma":[0.0000021727649,0.000011996979,0.0016761543,0.000001953622,0.0000039815395,0.000014713074,0.000009681165,0.99674076,0.00013016336,0.0013620652,0.00004272239,0.0000036449585],"about_ca_topic_score_codex":0.007641704,"about_ca_topic_score_gemma":0.005649428,"teacher_disagreement_score":0.007641704,"about_ca_system_score_codex":0.0011697103,"about_ca_system_score_gemma":0.0005588532,"threshold_uncertainty_score":0.015194476},"labels":[],"label_agreement":null},{"id":"W1991242370","doi":"10.1109/icst.2013.45","title":"Automated Detection of Test Fixture Strategies and Smells","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Fixture; Code refactoring; Code smell; Computer science; Test fixture; Test (biology); Code (set theory); Reliability engineering; Software engineering; Code coverage; Test Management Approach; Set (abstract data type); Engineering; Software quality; Software; Programming language; Software system; Software development; Software construction","score_opus":0.008291020522666195,"score_gpt":0.23610085659940022,"score_spread":0.22780983607673402,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1991242370","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6312799,0.0013242961,0.3395823,0.0005735476,0.00007744905,0.0004603048,0.001315486,0.02136363,0.004023146],"genre_scores_gemma":[0.8122084,0.0002762273,0.1839869,0.0001531501,0.000023275526,0.00021086939,0.0012516048,0.00076360727,0.0011259626],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916933,0.0019440933,0.0008104275,0.0013319979,0.0038314501,0.0003886655],"domain_scores_gemma":[0.90963995,0.050482295,0.019004593,0.008595759,0.010671789,0.0016056234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049882894,0.0013090763,0.0010011728,0.006735808,0.00046321316,0.0017816613,0.0016277357,0.001169761,0.00085119804],"category_scores_gemma":[0.049637567,0.0006997342,0.0006666296,0.0022758315,0.0006745247,0.001731838,0.0011421998,0.00088657666,0.0005273419],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075508666,0.00057017966,0.25996983,0.0011238416,0.00033149135,0.0030992476,0.003290041,0.020893358,0.14212522,0.0035577079,0.006045376,0.55823857],"study_design_scores_gemma":[0.00014766523,0.0016149464,0.23910198,0.00058250286,0.0003573069,0.006215305,0.0019106566,0.52165556,0.2038651,0.008275889,0.015848426,0.00042472134],"about_ca_topic_score_codex":0.001588703,"about_ca_topic_score_gemma":0.0030374755,"teacher_disagreement_score":0.006735808,"about_ca_system_score_codex":0.0008148798,"about_ca_system_score_gemma":0.0010554031,"threshold_uncertainty_score":0.026380897},"labels":[],"label_agreement":null},{"id":"W1991613282","doi":"10.1007/s10664-010-9150-8","title":"A field study of API learning obstacles","year":2010,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":352,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"McGill University; Microsoft Research","keywords":"Documentation; Computer science; Programmer; Presentation (obstetrics); Field (mathematics); World Wide Web; Software engineering; Software documentation; Internal documentation; Multimedia; Software development; Software; Programming language; Software development process","score_opus":0.017733432877838518,"score_gpt":0.2866929103016268,"score_spread":0.2689594774237883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1991613282","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9914458,0.000065053086,0.0017018118,0.00018560377,0.00001502878,0.00020139701,0.00010459795,0.00003871158,0.006242027],"genre_scores_gemma":[0.99280375,0.0001134372,0.001679601,0.0001119844,0.000012381041,0.00018487738,0.00017919597,0.000018074192,0.0048965234],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9978828,0.00092003145,0.00012410858,0.00032024586,0.00049273594,0.00026010018],"domain_scores_gemma":[0.9486915,0.034418553,0.0029169226,0.0030336631,0.007387905,0.0035515027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035845973,0.00039628908,0.00038244444,0.0024594048,0.0028134368,0.0014864256,0.0015115013,0.0011249878,0.0071304566],"category_scores_gemma":[0.026904583,0.00047759406,0.00029095422,0.001810293,0.0016359268,0.002756436,0.001628241,0.002297278,0.000967709],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0061581796,0.11965721,0.369503,0.0012044397,0.00009226475,0.0029006545,0.13451378,0.003056696,0.018081622,0.025637178,0.016447537,0.30274734],"study_design_scores_gemma":[0.0012575759,0.03595973,0.53708005,0.0008470285,0.00022087379,0.002551469,0.30097997,0.018330079,0.023863547,0.02203592,0.056562986,0.00031076482],"about_ca_topic_score_codex":0.006725766,"about_ca_topic_score_gemma":0.008847254,"teacher_disagreement_score":0.0071304566,"about_ca_system_score_codex":0.0017108708,"about_ca_system_score_gemma":0.003065803,"threshold_uncertainty_score":0.02385372},"labels":[],"label_agreement":null},{"id":"W1991829485","doi":"10.1109/compsacw.2013.11","title":"Opportunities for Clone Detection in Test Case Recommendation","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Programmer; clone (Java method); Unit testing; Computer science; Context (archaeology); Software engineering; Software maintenance; Knowledge base; Test case; Data mining; Test (biology); Software; Software development; Programming language; World Wide Web; Machine learning","score_opus":0.08035914689720865,"score_gpt":0.2903417199036316,"score_spread":0.20998257300642298,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1991829485","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17479706,0.0027051372,0.80771047,0.0019638764,0.000060472034,0.00045558493,0.00028265786,0.0062401807,0.0057844534],"genre_scores_gemma":[0.57754904,0.00066640123,0.41863066,0.00023451039,0.00007720878,0.00021082128,0.00076432433,0.00037449072,0.0014925086],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9854168,0.006903522,0.001048752,0.0018725183,0.004219914,0.00053844193],"domain_scores_gemma":[0.8192456,0.14806639,0.008358259,0.012245449,0.010840546,0.0012437752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011162172,0.0011493422,0.001268565,0.00682475,0.0009029631,0.0034478796,0.0027898243,0.0033260633,0.0022858481],"category_scores_gemma":[0.1097503,0.0011158941,0.0012168544,0.00302426,0.0009896947,0.0057893544,0.0017152799,0.0017241859,0.0010206777],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061128294,0.00057419733,0.06712485,0.00035970908,0.00019135047,0.0009717078,0.00200824,0.016808892,0.016415974,0.004906829,0.003065432,0.8869615],"study_design_scores_gemma":[0.00035568152,0.0016339308,0.045836125,0.0010305234,0.0008825575,0.006494514,0.002226383,0.79342103,0.07596359,0.037485234,0.034331657,0.00033874888],"about_ca_topic_score_codex":0.0053811595,"about_ca_topic_score_gemma":0.00886607,"teacher_disagreement_score":0.011162172,"about_ca_system_score_codex":0.0009278967,"about_ca_system_score_gemma":0.0016432843,"threshold_uncertainty_score":0.059031963},"labels":[],"label_agreement":null},{"id":"W1991867644","doi":"10.1007/s10664-015-9375-7","title":"Analyzing and automatically labelling the types of user issues that are raised in mobile app reviews","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":188,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Queen's University","funders":"","keywords":"App store; Computer science; World Wide Web; Download; Internet privacy; Mobile apps; Mobile device; Analytics; Notice; Data science","score_opus":0.052240163040090756,"score_gpt":0.31900221406896706,"score_spread":0.2667620510288763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1991867644","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92182285,0.008429823,0.043867704,0.002083048,0.00072946347,0.0010159647,0.006626593,0.003083073,0.012341457],"genre_scores_gemma":[0.90672195,0.0024087566,0.073671125,0.0005298289,0.00040555844,0.00060249533,0.008159434,0.00039396717,0.0071069133],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.98891246,0.0035201486,0.0013476085,0.0011936752,0.0046751774,0.00035094388],"domain_scores_gemma":[0.8451639,0.108508475,0.0188138,0.0036064289,0.022743845,0.0011636401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006229083,0.00080312986,0.0007328273,0.012319657,0.001039726,0.0027660078,0.0007881423,0.0014499133,0.0010666024],"category_scores_gemma":[0.08608788,0.0005188828,0.0006743261,0.0041393484,0.0003684386,0.002666769,0.0013023436,0.0010668656,0.0011805536],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011599872,0.000435178,0.360364,0.0075349472,0.00046871678,0.0028703157,0.012884311,0.0013878444,0.056433152,0.0029637518,0.043104563,0.51039326],"study_design_scores_gemma":[0.000116610354,0.0011124874,0.6913583,0.0031181115,0.0014892172,0.008753801,0.013747976,0.07616106,0.06316012,0.0056732856,0.13487162,0.00043739166],"about_ca_topic_score_codex":0.0037181955,"about_ca_topic_score_gemma":0.01187222,"teacher_disagreement_score":0.012319657,"about_ca_system_score_codex":0.00069938286,"about_ca_system_score_gemma":0.0019277618,"threshold_uncertainty_score":0.03294295},"labels":[],"label_agreement":null},{"id":"W1991987529","doi":"10.5555/381473.381601","title":"How to do inspections when there is no time","year":2001,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Software inspection; Computer science; Software; Statistical process control; Quality (philosophy); Software quality; Software development; Software engineering; Process (computing)","score_opus":0.02655269819784756,"score_gpt":0.2669592391930008,"score_spread":0.24040654099515324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1991987529","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011282475,0.0030225003,0.8876302,0.015623863,0.0026995097,0.0010205267,0.0006288932,0.023786498,0.054305594],"genre_scores_gemma":[0.06898747,0.003271334,0.8807529,0.0035300318,0.0010368578,0.00073522236,0.0010450853,0.0033577268,0.03728333],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99491394,0.0010808667,0.00031361825,0.00087721786,0.0024092328,0.0004051797],"domain_scores_gemma":[0.98903614,0.0027842105,0.00095042755,0.0032476888,0.0032196934,0.0007618332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043055997,0.001538796,0.001496922,0.0019277489,0.0017886584,0.0044426695,0.0022150371,0.002592985,0.02674942],"category_scores_gemma":[0.030435344,0.001164169,0.00085124397,0.001123489,0.0016268604,0.007897033,0.001418722,0.003064483,0.033729438],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001578872,0.00021958169,0.0017929026,0.0005014387,0.00005410757,0.0002499045,0.00096000946,0.0033047914,0.01506938,0.013900536,0.09800491,0.8657845],"study_design_scores_gemma":[0.00019568573,0.0004300454,0.0072633415,0.001272827,0.00016291664,0.0024369308,0.0031838447,0.03663812,0.035345577,0.16849612,0.74410075,0.00047387352],"about_ca_topic_score_codex":0.002203825,"about_ca_topic_score_gemma":0.003115728,"teacher_disagreement_score":0.02674942,"about_ca_system_score_codex":0.0005776299,"about_ca_system_score_gemma":0.0019997063,"threshold_uncertainty_score":0.089485765},"labels":[],"label_agreement":null},{"id":"W1992312374","doi":"10.1016/s0164-1212(01)00049-8","title":"Using self-organizing maps to analyze object-oriented software measures","year":2001,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software visualization; Computer science; Software; Software sizing; Visualization; Software construction; Software metric; Software development; Domain (mathematical analysis); Software system; Software analytics; Software deployment; Java; Software engineering; Software framework; Data mining; Human–computer interaction; Programming language","score_opus":0.02772940442969858,"score_gpt":0.2707078016126779,"score_spread":0.2429783971829793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1992312374","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41204146,0.00046034894,0.5828732,0.00012234655,0.000058115464,0.0001799946,0.00029010413,0.0014599086,0.0025145586],"genre_scores_gemma":[0.7992775,0.00012008167,0.19949889,0.0000133076255,0.00003321658,0.000120591896,0.00039604338,0.000107217376,0.00043318517],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99891317,0.00035877916,0.000089067085,0.00012892274,0.00042519683,0.0000848386],"domain_scores_gemma":[0.99287814,0.004615242,0.00072051433,0.00038968856,0.0012326628,0.00016389159],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025084275,0.000487989,0.00045210507,0.008920738,0.00057499495,0.0017671898,0.00049438595,0.00040451385,0.000724879],"category_scores_gemma":[0.0093727745,0.00024599163,0.00066795043,0.0038831518,0.00074256945,0.0018253007,0.0006718449,0.00047182967,0.00019054791],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006300671,0.0005534801,0.09537654,0.0005424276,0.0006345826,0.00034884622,0.0029151184,0.08545132,0.023238162,0.04477486,0.0031556638,0.742379],"study_design_scores_gemma":[0.0000551227,0.0002861549,0.09116609,0.00005477871,0.00020035461,0.00029018757,0.0017230626,0.8265734,0.013489449,0.06289422,0.003153055,0.00011415667],"about_ca_topic_score_codex":0.0030607178,"about_ca_topic_score_gemma":0.002496741,"teacher_disagreement_score":0.008920738,"about_ca_system_score_codex":0.00061562617,"about_ca_system_score_gemma":0.0006692816,"threshold_uncertainty_score":0.013266027},"labels":[],"label_agreement":null},{"id":"W1992377888","doi":"10.1109/icmla.2012.138","title":"Estimating Software Effort Using an ANN Model Based on Use Case Points","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Use Case Points; Artificial neural network; Computer science; Software; Regression analysis; Linear regression; Data modeling; Data mining; Machine learning; Artificial intelligence; Software development; Software development process","score_opus":0.07663196190569285,"score_gpt":0.32466701358021305,"score_spread":0.2480350516745202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1992377888","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33505693,0.00038372132,0.6594756,0.00024052181,0.000043519427,0.00013273554,0.0003808763,0.0010983694,0.0031876073],"genre_scores_gemma":[0.89532405,0.0002769449,0.10163565,0.000036482907,0.000022996712,0.00019514294,0.00062034756,0.00004109558,0.0018472996],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99888426,0.0003802505,0.000096873344,0.0002196071,0.0003517413,0.00006724724],"domain_scores_gemma":[0.9951746,0.0031522526,0.0005579748,0.00023257654,0.00081026135,0.00007234753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016118989,0.0009900781,0.0006530001,0.0019046149,0.00021953393,0.00095792575,0.0012693633,0.0009854949,0.0012334466],"category_scores_gemma":[0.008559807,0.000506535,0.0006481774,0.0013921923,0.0002495364,0.0018340695,0.0004471662,0.00075121043,0.00046435025],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012348131,0.00015266998,0.016502148,0.00009108224,0.00012446642,0.00012387798,0.000068255315,0.9049419,0.0016145452,0.00079691893,0.00039528153,0.07506527],"study_design_scores_gemma":[0.00000226204,0.000019303627,0.0016191656,0.000007824857,0.00000814411,0.000015265286,0.000006791809,0.9974802,0.00036698286,0.0003669825,0.00010178172,0.0000052319465],"about_ca_topic_score_codex":0.0058463286,"about_ca_topic_score_gemma":0.006312782,"teacher_disagreement_score":0.0058463286,"about_ca_system_score_codex":0.0006892501,"about_ca_system_score_gemma":0.0004187089,"threshold_uncertainty_score":0.011624575},"labels":[],"label_agreement":null},{"id":"W1992706013","doi":"10.1109/icsm.2011.6080819","title":"Code convention adherence in evolving software","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"University of Alberta","keywords":"Computer science; Software engineering; Programming language; Code (set theory); Convention; Software; Political science","score_opus":0.055244442357530024,"score_gpt":0.2756993334609918,"score_spread":0.22045489110346178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1992706013","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9401063,0.00025867583,0.055183645,0.0003310351,0.00002630985,0.00015499802,0.0001348444,0.00027152276,0.0035326811],"genre_scores_gemma":[0.97605604,0.00010885886,0.022888413,0.00005457213,0.000017820312,0.00009699806,0.00022920313,0.000068297515,0.00047984178],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.96850026,0.010446278,0.005004351,0.0027099163,0.012506148,0.00083306635],"domain_scores_gemma":[0.71317446,0.10086577,0.10410973,0.036742102,0.041445825,0.003662232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016542792,0.00028582054,0.0003526285,0.00535322,0.0016659049,0.0027444763,0.0011715887,0.0010967782,0.000602359],"category_scores_gemma":[0.20102993,0.00032776612,0.00044577112,0.004333808,0.0031644579,0.0041594286,0.0030577192,0.0016531794,0.00013377785],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003474613,0.00035091458,0.6743411,0.00036535168,0.000151209,0.00096195313,0.03267051,0.013355235,0.018362923,0.028664254,0.0011689938,0.22926009],"study_design_scores_gemma":[0.000047323054,0.0011173445,0.8030487,0.00039888656,0.00020836663,0.0027219853,0.01258859,0.08937401,0.024365481,0.049397234,0.016450794,0.00028124737],"about_ca_topic_score_codex":0.004475092,"about_ca_topic_score_gemma":0.0038346872,"teacher_disagreement_score":0.016542792,"about_ca_system_score_codex":0.0017135858,"about_ca_system_score_gemma":0.0018111322,"threshold_uncertainty_score":0.0874877},"labels":[],"label_agreement":null},{"id":"W1992943987","doi":"10.1145/1140123.1140237","title":"Finite automata models for CS problem with binary semaphore","year":2006,"lang":"en","type":"article","venue":"ACM SIGCSE Bulletin","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nipissing University","funders":"","keywords":"Semaphore; Computer science; Synchronization (alternating current); Automaton; Simplicity; Deterministic finite automaton; Finite-state machine; Visualization; Theoretical computer science; Binary number; ω-automaton; Quantum finite automata; Algorithm; Automata theory; Parallel computing; Programming language; Mathematics; Arithmetic; Artificial intelligence; Channel (broadcasting)","score_opus":0.01652126303611729,"score_gpt":0.23239711755673698,"score_spread":0.2158758545206197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1992943987","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022382723,0.0009624059,0.95642906,0.00095640327,0.00016110466,0.00014758944,0.00036897152,0.00037446382,0.018217284],"genre_scores_gemma":[0.59650934,0.0019481045,0.36426783,0.00038599627,0.0003989516,0.0008170688,0.0010096527,0.00023920677,0.034423947],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991455,0.00028046398,0.000061259576,0.00020882764,0.00019056251,0.000113392074],"domain_scores_gemma":[0.99779725,0.0015324088,0.00018999273,0.00019260684,0.00018664607,0.00010116328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00088466157,0.00087652024,0.00060793944,0.00075748307,0.0009790933,0.0020441124,0.0013768827,0.0016340151,0.006490615],"category_scores_gemma":[0.0037045504,0.0003710696,0.0015380599,0.00086990226,0.0013525815,0.0026527871,0.001178644,0.0022578677,0.0009022248],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058467358,0.000069897964,0.0005475119,0.0001481225,0.000026411295,0.00029457707,0.00025407743,0.1687563,0.0015569276,0.8142946,0.0026281723,0.011364841],"study_design_scores_gemma":[0.000018204404,0.00003530009,0.00012838184,0.000040825253,0.000024392919,0.00012611189,0.0000917172,0.5071685,0.0006270992,0.48360187,0.008120314,0.000017359718],"about_ca_topic_score_codex":0.0055445037,"about_ca_topic_score_gemma":0.0066098864,"teacher_disagreement_score":0.006490615,"about_ca_system_score_codex":0.0018163959,"about_ca_system_score_gemma":0.0021403397,"threshold_uncertainty_score":0.021713257},"labels":[],"label_agreement":null},{"id":"W1993507155","doi":"10.1145/2001858.2001996","title":"Software clustering by example","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Cluster analysis; Computer science; Graph; Software; Clustering coefficient; Theoretical computer science; Data mining; Programming language; Artificial intelligence","score_opus":0.050462155814105146,"score_gpt":0.24698800405891333,"score_spread":0.19652584824480818,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1993507155","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020432709,0.00020404131,0.96761113,0.00040785444,0.000037491758,0.000083321764,0.00012371555,0.0003112259,0.010788522],"genre_scores_gemma":[0.2925731,0.00042122856,0.69543934,0.00018597327,0.000045720764,0.00023089479,0.00051547185,0.00025032493,0.010337896],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986148,0.0005059736,0.00007123204,0.00028391287,0.00038742617,0.0001367768],"domain_scores_gemma":[0.9979163,0.0009828216,0.00014679207,0.0004916958,0.00037401108,0.00008842412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010143719,0.0007875819,0.00050594436,0.0012361994,0.0008092403,0.0009361588,0.0017030711,0.0014042786,0.006118735],"category_scores_gemma":[0.0049908957,0.0002842268,0.0011184631,0.0018799773,0.0013828415,0.0022677267,0.0018773566,0.0016447784,0.0010751847],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009057143,0.00010610475,0.0016194008,0.0002528899,0.000052346535,0.0003715977,0.00053138786,0.3822953,0.0038894948,0.5269968,0.007388912,0.07640519],"study_design_scores_gemma":[0.000026732507,0.000058257898,0.00041673315,0.00004366337,0.000029961067,0.000368703,0.00016548621,0.6764522,0.0035849672,0.2841429,0.034686036,0.000024339108],"about_ca_topic_score_codex":0.0023026417,"about_ca_topic_score_gemma":0.0039236713,"teacher_disagreement_score":0.006118735,"about_ca_system_score_codex":0.00074481004,"about_ca_system_score_gemma":0.00044110717,"threshold_uncertainty_score":0.020469248},"labels":[],"label_agreement":null},{"id":"W1993673250","doi":"10.1109/saner.2015.7081830","title":"Threshold-free code clone detection for a large-scale heterogeneous Java repository","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Java; Granularity; Computer science; Benchmark (surveying); Source code; Code (set theory); Data mining; Software; Variety (cybernetics); Algorithm; Artificial intelligence; Operating system; Programming language; Biology","score_opus":0.02652208983493504,"score_gpt":0.2673020073773731,"score_spread":0.24077991754243805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1993673250","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42347682,0.00045084587,0.56406045,0.00019357317,0.000044049924,0.00019546422,0.00034389796,0.010611673,0.00062315393],"genre_scores_gemma":[0.6945908,0.000091852955,0.3034626,0.00006962573,0.000017219038,0.00012183995,0.0008614311,0.00022859396,0.00055597204],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99555594,0.00057952496,0.0004552363,0.0013950265,0.0017531442,0.00026113493],"domain_scores_gemma":[0.981768,0.006210736,0.003456235,0.0033034112,0.004586378,0.0006752803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028345655,0.0008390485,0.0013061273,0.0051964233,0.0011318824,0.0015501559,0.0025999078,0.0015955658,0.00030797906],"category_scores_gemma":[0.019684987,0.00044420565,0.001144004,0.0036069774,0.00074154866,0.002520254,0.0016861791,0.0010169181,0.00032913292],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060187414,0.0007166192,0.14231461,0.0004695969,0.00040494293,0.0010745418,0.0011003828,0.09451389,0.12128725,0.0025390722,0.004638421,0.6303388],"study_design_scores_gemma":[0.000035561698,0.00022025057,0.027274545,0.000020914256,0.00007870889,0.000919533,0.00025767833,0.91784793,0.048338857,0.0035269663,0.0014186035,0.00006051787],"about_ca_topic_score_codex":0.005267264,"about_ca_topic_score_gemma":0.00695559,"teacher_disagreement_score":0.005267264,"about_ca_system_score_codex":0.0011402316,"about_ca_system_score_gemma":0.0013310005,"threshold_uncertainty_score":0.014990747},"labels":[],"label_agreement":null},{"id":"W1993743759","doi":"10.1016/j.jss.2013.12.037","title":"Cooperation, collaboration and pair-programming: Field studies on backup behavior","year":2014,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Backup; Teamwork; Field (mathematics); Computer science; Pair programming; Salient; Work (physics); Artificial intelligence; Engineering; Programming language; Software development; Database; Software; Mathematics; Management","score_opus":0.02510894068006842,"score_gpt":0.2978236989630619,"score_spread":0.27271475828299346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1993743759","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9985018,0.000028622286,0.00028023814,0.00003923621,0.0000035189141,0.00003348218,0.00002087732,0.0000045491993,0.0010875965],"genre_scores_gemma":[0.99866927,0.000045306402,0.0004230777,0.00003997471,0.0000043505074,0.000063743944,0.000038559665,0.0000053759773,0.000710233],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.996128,0.0028127346,0.00014616444,0.00040989748,0.00029557248,0.00020762156],"domain_scores_gemma":[0.9150917,0.061796814,0.005951387,0.0069273664,0.0059257145,0.0043069893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007254655,0.0004248518,0.00050179474,0.0018262324,0.0034952909,0.0011843104,0.0015184241,0.0011340156,0.003468371],"category_scores_gemma":[0.034993943,0.00054915744,0.0002513454,0.0011725711,0.0033162108,0.0020694674,0.0017951988,0.0016874038,0.0004244935],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0116475355,0.047767162,0.5936081,0.00076732406,0.00021385065,0.0014492266,0.20137146,0.0030015847,0.018813448,0.009496288,0.002949202,0.10891484],"study_design_scores_gemma":[0.0011147683,0.023875177,0.687985,0.0004452735,0.00017049075,0.0023107051,0.23469816,0.009340645,0.012830266,0.015093235,0.011870076,0.00026614804],"about_ca_topic_score_codex":0.0036239745,"about_ca_topic_score_gemma":0.004687788,"teacher_disagreement_score":0.007254655,"about_ca_system_score_codex":0.0009273165,"about_ca_system_score_gemma":0.0014530988,"threshold_uncertainty_score":0.038366795},"labels":[],"label_agreement":null},{"id":"W1993807037","doi":"10.1016/j.scico.2009.10.007","title":"Rigi—An environment for software reverse engineering, exploration, visualization, and redocumentation","year":2009,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Center for Advanced Study, University of Illinois at Urbana-Champaign; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Reverse engineering; Software engineering; Visualization; Interoperability; Scripting language; Extensibility; Parsing; Software; Process (computing); Programming language; World Wide Web; Artificial intelligence","score_opus":0.016308934371584573,"score_gpt":0.28788989181511715,"score_spread":0.27158095744353256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1993807037","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025920055,0.0004425075,0.5278517,0.00022784983,0.00015109686,0.00028628256,0.0036845976,0.4558502,0.008913747],"genre_scores_gemma":[0.04957659,0.0012383126,0.84840363,0.0004965256,0.0001302922,0.0010185362,0.016583463,0.06724218,0.015310484],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99744105,0.00067833735,0.00017814181,0.0004119161,0.0010923641,0.00019820916],"domain_scores_gemma":[0.99511343,0.0023978925,0.00024123356,0.0016144775,0.00037285534,0.00026014558],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00384071,0.0024328528,0.0019324734,0.002224318,0.00086423557,0.0036440734,0.005661873,0.001617777,0.03238906],"category_scores_gemma":[0.008893183,0.0017753307,0.002466921,0.0014168688,0.0010300144,0.0044647628,0.0055571673,0.003945358,0.013074082],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023873374,0.0006854179,0.0025934877,0.0024119269,0.00057754817,0.0009843857,0.0016087212,0.02192651,0.030425996,0.06405771,0.36081034,0.5115306],"study_design_scores_gemma":[0.0012043414,0.00041517674,0.0019868594,0.0005016761,0.0002607562,0.0010297948,0.00025543084,0.27431232,0.058930036,0.06858795,0.59201354,0.0005021323],"about_ca_topic_score_codex":0.0023504035,"about_ca_topic_score_gemma":0.0040335353,"teacher_disagreement_score":0.03238906,"about_ca_system_score_codex":0.0006067924,"about_ca_system_score_gemma":0.0018403339,"threshold_uncertainty_score":0.108352244},"labels":[],"label_agreement":null},{"id":"W1994059284","doi":"10.1109/dapse.2013.6603803","title":"Extracting artifact lifecycle models from metadata history","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artifact (error); System lifecycle; Software engineering; Metadata; Software development process; Process (computing); Goal-Driven Software Development Process; Software development; Application lifecycle management; Iterative and incremental development; Set (abstract data type); Software; Systems engineering; Artificial intelligence; Programming language; Engineering; World Wide Web","score_opus":0.061600614161988365,"score_gpt":0.24467345658097497,"score_spread":0.1830728424189866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994059284","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08910781,0.0021351527,0.859362,0.0015374557,0.00012313454,0.0011922504,0.019122964,0.008244502,0.0191747],"genre_scores_gemma":[0.3150481,0.0026082164,0.6331717,0.00011543647,0.000065936125,0.00097412855,0.041732017,0.000908549,0.0053759078],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981792,0.00042700954,0.0003262961,0.0002443533,0.0007039459,0.00011926285],"domain_scores_gemma":[0.9862,0.0058276504,0.0018020184,0.0028413713,0.0029692359,0.0003597706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033497468,0.000982612,0.0006220215,0.013318667,0.00097470527,0.0037096082,0.0012705935,0.0011452956,0.0026411565],"category_scores_gemma":[0.021667186,0.0010653416,0.0015071098,0.007891148,0.00052409555,0.006828798,0.0018824448,0.0011338898,0.0021977487],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028474073,0.00040789097,0.058431607,0.0017725173,0.00017141017,0.002203402,0.005354824,0.03791511,0.007054808,0.0773882,0.020906344,0.7881091],"study_design_scores_gemma":[0.00011691773,0.00037691402,0.04315432,0.0026136534,0.00063687185,0.0021416927,0.0070067635,0.4688372,0.023497852,0.15435937,0.29686773,0.00039071083],"about_ca_topic_score_codex":0.016163856,"about_ca_topic_score_gemma":0.022922887,"teacher_disagreement_score":0.016163856,"about_ca_system_score_codex":0.0020281312,"about_ca_system_score_gemma":0.0056206496,"threshold_uncertainty_score":0.03213954},"labels":[],"label_agreement":null},{"id":"W1994243621","doi":"10.1002/smr.327","title":"Supporting the analysis of clones in software systems","year":2006,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":74,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cloning (programming); Computer science; clone (Java method); Software system; Software; Software engineering; Source code; Software evolution; Code (set theory); Software development; Programming language; Software construction; Biology; Genetics; Set (abstract data type)","score_opus":0.029829150576390785,"score_gpt":0.34921757749210136,"score_spread":0.31938842691571057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994243621","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58692,0.00043637175,0.40155214,0.00061803265,0.000022040795,0.00029325546,0.00037786478,0.0067317015,0.0030486067],"genre_scores_gemma":[0.80913085,0.00009475975,0.18959066,0.000044451775,0.000014708969,0.00008184389,0.00041817245,0.00021989335,0.00040476155],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9891089,0.0044507137,0.0008594226,0.0011118118,0.0040545315,0.00041457528],"domain_scores_gemma":[0.8606563,0.10209808,0.012274185,0.013833752,0.010150979,0.0009867711],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070882253,0.0006708995,0.00079148816,0.0049152216,0.0009513675,0.002347547,0.0016390238,0.0011162099,0.0012625915],"category_scores_gemma":[0.08424467,0.00057770096,0.0006614503,0.003990873,0.0016003203,0.0043759705,0.0024032074,0.0009531248,0.0003555547],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000646673,0.0005358861,0.19169562,0.0011617258,0.00024441484,0.0028540145,0.015778577,0.05518237,0.07453212,0.025684433,0.003211097,0.6284731],"study_design_scores_gemma":[0.00008954038,0.0005096328,0.04799316,0.00030710973,0.00017851747,0.0018899543,0.0034553567,0.76503724,0.11496064,0.05641105,0.009008711,0.00015907788],"about_ca_topic_score_codex":0.0026690403,"about_ca_topic_score_gemma":0.002147793,"teacher_disagreement_score":0.0070882253,"about_ca_system_score_codex":0.00092566886,"about_ca_system_score_gemma":0.001343122,"threshold_uncertainty_score":0.037486613},"labels":[],"label_agreement":null},{"id":"W1994248747","doi":"10.1145/2597073.2597075","title":"An empirical study of just-in-time defect prediction using cross-project models","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":178,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Leverage (statistics); Predictive modelling; Computer science; Context (archaeology); Project management; Data science; Machine learning; Artificial intelligence; Systems engineering; Engineering","score_opus":0.07860207803743692,"score_gpt":0.38833946920076146,"score_spread":0.30973739116332455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994248747","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96871644,0.0008189816,0.027639382,0.00047354153,0.000045314744,0.0000838223,0.0006524219,0.0002799608,0.0012901517],"genre_scores_gemma":[0.9878748,0.00020274104,0.00957046,0.000065880355,0.000023532428,0.000068891924,0.0017369846,0.000056523113,0.00040011707],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98816395,0.006687156,0.00078130426,0.0024473537,0.0013831401,0.0005369393],"domain_scores_gemma":[0.7996609,0.16466734,0.009669901,0.014664115,0.0090376325,0.002300119],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03509132,0.0016569196,0.001270124,0.0023263034,0.00085849233,0.0022336035,0.0024870655,0.0017450382,0.0012466836],"category_scores_gemma":[0.096422076,0.0007895266,0.0010019796,0.0027708074,0.0011386399,0.006513715,0.002179768,0.0039459425,0.00068205746],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009748382,0.0016549845,0.5108214,0.0003896258,0.00069679774,0.00040091327,0.0013968395,0.38044024,0.0009983219,0.00325725,0.0050068586,0.09396184],"study_design_scores_gemma":[0.0000427224,0.0006357153,0.044692244,0.00007252903,0.00008388429,0.00021286627,0.0004192683,0.9482621,0.0009109836,0.0035192661,0.0010971026,0.000051365492],"about_ca_topic_score_codex":0.008642159,"about_ca_topic_score_gemma":0.008380101,"teacher_disagreement_score":0.03509132,"about_ca_system_score_codex":0.0013067188,"about_ca_system_score_gemma":0.0010477698,"threshold_uncertainty_score":0.18558288},"labels":[],"label_agreement":null},{"id":"W1994492702","doi":"10.1145/1022494.1022542","title":"Continuous evolutionary one-step-ahead testing","year":2004,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; University of Waterloo","funders":"","keywords":"Vendor; Computer science; Software engineering; Software release life cycle; Software; Software development; Backporting; Reliability engineering; Software reliability testing; Software construction; Engineering; Operating system; Business","score_opus":0.029425337035696504,"score_gpt":0.24813374841376354,"score_spread":0.21870841137806704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994492702","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13339412,0.0004828022,0.8405125,0.00062531134,0.00017144308,0.0005690264,0.0001717414,0.004812943,0.019260209],"genre_scores_gemma":[0.7139198,0.00026131986,0.27466875,0.00031532103,0.00006173701,0.00034481887,0.00035764152,0.00018283194,0.009887757],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9961746,0.0010230407,0.00011346449,0.00058608636,0.0018169283,0.0002859115],"domain_scores_gemma":[0.988399,0.0051868395,0.00069605844,0.0026230132,0.0024311503,0.00066393946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034157773,0.000986926,0.0007246662,0.0009231938,0.00072752364,0.001395299,0.0032728268,0.0013270216,0.0057196743],"category_scores_gemma":[0.012805851,0.0004345545,0.00062237465,0.00065805775,0.0011523477,0.0014867584,0.0022092306,0.0012575076,0.0014096879],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008297842,0.0015217355,0.0150877265,0.000521215,0.00015757594,0.0010248143,0.0010132372,0.11831382,0.042032134,0.029950036,0.0052720658,0.7842758],"study_design_scores_gemma":[0.00034844154,0.003659206,0.011425345,0.0002063078,0.00018631053,0.0024417017,0.0004286824,0.8794731,0.03673367,0.036300305,0.028639426,0.0001575435],"about_ca_topic_score_codex":0.0011492684,"about_ca_topic_score_gemma":0.0012670832,"teacher_disagreement_score":0.0057196743,"about_ca_system_score_codex":0.0005577057,"about_ca_system_score_gemma":0.0013803219,"threshold_uncertainty_score":0.019134164},"labels":[],"label_agreement":null},{"id":"W1994493193","doi":"10.5555/2819009.2819026","title":"Online defect prediction for imbalanced data","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":156,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Commit; Process (computing); Machine learning; Reliability (semiconductor); Data mining; Statistical classification; Predictive modelling; Resampling; Software bug; Artificial intelligence; Change detection; Software; Database","score_opus":0.1116123255897715,"score_gpt":0.34237113988862944,"score_spread":0.23075881429885794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994493193","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8438423,0.0011238939,0.14417209,0.0018471617,0.0004217091,0.00012543808,0.0027759902,0.0041024266,0.0015889308],"genre_scores_gemma":[0.96481997,0.0001066099,0.031438634,0.00013713063,0.0001843247,0.00005166101,0.0028247153,0.00006807889,0.00036884827],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950859,0.0018960001,0.00034191267,0.001100769,0.0012329563,0.0003424346],"domain_scores_gemma":[0.9473374,0.031878773,0.0053814915,0.00810714,0.006234489,0.0010606516],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007483789,0.0011071353,0.0012626443,0.0052916827,0.000661167,0.0013610545,0.0014264596,0.001223704,0.00073020963],"category_scores_gemma":[0.029601235,0.0003332648,0.0006673695,0.003029001,0.0006228384,0.0023380148,0.0011862449,0.0022155726,0.0007584977],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012074017,0.0013400462,0.3811529,0.00017316885,0.00043594223,0.0005158589,0.00027241337,0.23896307,0.0060355556,0.0012253769,0.017520523,0.35115778],"study_design_scores_gemma":[0.000017995913,0.00008800172,0.02112061,0.000014843848,0.000017546528,0.000080020945,0.000069247006,0.97355086,0.0021759532,0.002248308,0.0006018548,0.000014696497],"about_ca_topic_score_codex":0.0039674295,"about_ca_topic_score_gemma":0.0031063522,"teacher_disagreement_score":0.007483789,"about_ca_system_score_codex":0.0009294336,"about_ca_system_score_gemma":0.0006172639,"threshold_uncertainty_score":0.039578557},"labels":[],"label_agreement":null},{"id":"W1994504150","doi":"10.5539/cis.v2n3p64","title":"Positive Affects Inducer on Software Quality","year":2009,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Universiti Utara Malaysia","keywords":"Agile software development; Computer science; Extreme programming; Metric (unit); Quality (philosophy); Software; Contentment; Empirical research; Affect (linguistics); Extreme programming practices; Software development; Software engineering; Psychology; Statistics; Software development process; Operations management; Social psychology; Mathematics","score_opus":0.019249059278220505,"score_gpt":0.30061715558978236,"score_spread":0.28136809631156184,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994504150","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9608211,0.00031797192,0.003502802,0.00027579756,0.00003948055,0.000041055937,0.0000371097,0.000050021594,0.034914576],"genre_scores_gemma":[0.99865496,0.00007882339,0.00049229845,0.00003890526,0.000013913128,0.000012949405,0.000008212669,0.0000061296832,0.00069383177],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9969308,0.0014781699,0.0001122259,0.00017224331,0.001073998,0.00023261043],"domain_scores_gemma":[0.98126465,0.010766461,0.0032119271,0.0007889526,0.002130565,0.0018374432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016806393,0.0003614081,0.00018932366,0.0005716714,0.000771638,0.001456124,0.00016395094,0.0003298331,0.0027518324],"category_scores_gemma":[0.0129709495,0.00014066389,0.00026978712,0.00029554713,0.00090190297,0.0004867421,0.0015597014,0.0006371323,0.0002202623],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014365102,0.0012377258,0.5756605,0.0011325502,0.00035458407,0.0019053477,0.03580655,0.0032328877,0.05122251,0.020733858,0.0023571507,0.30491975],"study_design_scores_gemma":[0.000034202007,0.0014809844,0.9589192,0.0001269751,0.00022956285,0.0006951986,0.009171609,0.0023243173,0.009716049,0.005828943,0.011389727,0.0000832923],"about_ca_topic_score_codex":0.0003644514,"about_ca_topic_score_gemma":0.0004304192,"teacher_disagreement_score":0.0027518324,"about_ca_system_score_codex":0.0005692209,"about_ca_system_score_gemma":0.00029891304,"threshold_uncertainty_score":0.009205759},"labels":[],"label_agreement":null},{"id":"W1994514428","doi":"10.1049/iet-sen:20080010","title":"Reducing the use of nullable types through non-null by default and monotonic non-null","year":2008,"lang":"en","type":"article","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Programming language; Null (SQL); Java; Computer science; Software engineering; Java Modeling Language; Empirical research; Semantics (computer science); Java annotation; Database; Mathematics; Real time Java","score_opus":0.03589992865678084,"score_gpt":0.2562885052394724,"score_spread":0.22038857658269154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994514428","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052897837,0.00083849847,0.92167455,0.0030122362,0.00031514303,0.00019082299,0.0005071269,0.0076003084,0.012963433],"genre_scores_gemma":[0.31469226,0.00056879874,0.6679121,0.0017035777,0.00012612548,0.00034991864,0.0008180916,0.004503503,0.009325631],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.95710945,0.017888933,0.0057127955,0.0041573006,0.0136206355,0.0015108564],"domain_scores_gemma":[0.74239844,0.117248975,0.017456872,0.090948835,0.030064272,0.0018826667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.044847444,0.0011223663,0.0011803726,0.0044520386,0.0019877458,0.0067575336,0.006492913,0.0026125605,0.0054016355],"category_scores_gemma":[0.14447457,0.0017165976,0.0017630173,0.0035172761,0.0055709025,0.015688187,0.010317294,0.004979202,0.0022121377],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057549024,0.00033199583,0.052956097,0.0017908908,0.00013044324,0.0009691381,0.016012317,0.0056882245,0.024311958,0.47825304,0.011650258,0.4073302],"study_design_scores_gemma":[0.00014211258,0.00059476524,0.016560858,0.0028224983,0.00046161254,0.00548196,0.005586093,0.05357638,0.09125863,0.31432414,0.50848794,0.0007030573],"about_ca_topic_score_codex":0.0048168884,"about_ca_topic_score_gemma":0.008721945,"teacher_disagreement_score":0.044847444,"about_ca_system_score_codex":0.0026382152,"about_ca_system_score_gemma":0.006508856,"threshold_uncertainty_score":0.2371788},"labels":[],"label_agreement":null},{"id":"W1994598608","doi":"10.1145/2597073.2597121","title":"An insight into the pull requests of GitHub","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Order (exchange); Request for proposal; Key (lock); Information system; Base (topology)","score_opus":0.012574602474931963,"score_gpt":0.26861960825503306,"score_spread":0.2560450057801011,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994598608","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9734856,0.00032825128,0.013473457,0.002192885,0.000048165748,0.000119605385,0.00047810594,0.000553424,0.009320525],"genre_scores_gemma":[0.9809934,0.00035759743,0.009668534,0.0005072466,0.000069413196,0.00022849247,0.00074592995,0.0007519753,0.006677344],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9924948,0.0035774508,0.00056782184,0.00060574326,0.0021555435,0.0005985517],"domain_scores_gemma":[0.90679073,0.07422791,0.0074444376,0.002419214,0.0076354034,0.001482284],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0086706355,0.00081427663,0.00034795117,0.0058864434,0.0027868873,0.0030604913,0.0007569469,0.0012230747,0.0023376679],"category_scores_gemma":[0.048932906,0.00057566696,0.000257606,0.0033377714,0.0018637113,0.0051814234,0.0026695838,0.0016066313,0.00086188194],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048102823,0.00022935215,0.1912035,0.0010225612,0.00006098319,0.0048684212,0.6423076,0.0007934082,0.041350827,0.006340771,0.0097839935,0.10155756],"study_design_scores_gemma":[0.000042767202,0.0005226619,0.44202873,0.0006649477,0.0000967329,0.004745399,0.38746083,0.008803786,0.02177928,0.008217642,0.12524652,0.00039073266],"about_ca_topic_score_codex":0.0028167376,"about_ca_topic_score_gemma":0.0061701355,"teacher_disagreement_score":0.9913294,"about_ca_system_score_codex":0.0014794413,"about_ca_system_score_gemma":0.0015586807,"threshold_uncertainty_score":0.045855284},"labels":[],"label_agreement":null},{"id":"W1994641861","doi":"10.1109/icas.2010.24","title":"A Decentralized Evolutional Approach to Handle Schedule Execution in Software Projects","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Schedule; Computer science; Adaptation (eye); Software; Scheduling (production processes); Software project management; Task (project management); Software engineering; Project management; Distributed computing; Real-time computing; Software development; Systems engineering; Software construction; Engineering; Operations management; Programming language; Operating system","score_opus":0.019210604527091075,"score_gpt":0.26189680350052885,"score_spread":0.24268619897343777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994641861","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014432441,0.00005265368,0.98372674,0.00010039193,0.000015938786,0.00004668409,0.000020259651,0.0002929039,0.00131194],"genre_scores_gemma":[0.44289687,0.00012228476,0.554147,0.000055683806,0.00004222503,0.00018646919,0.00010588638,0.00011014937,0.002333433],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991062,0.0003648007,0.000050899667,0.00015775194,0.00026104468,0.0000593013],"domain_scores_gemma":[0.9989053,0.00042657257,0.00016017095,0.00023720205,0.00018698738,0.00008379262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015949187,0.00031050097,0.00042636695,0.00055909215,0.00073977874,0.00073046854,0.00135479,0.00049301697,0.0010795668],"category_scores_gemma":[0.003922809,0.00034800376,0.0004445333,0.00072532846,0.00063839875,0.0010731458,0.0011562681,0.00080618425,0.00016690356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000090151596,0.00013566684,0.0020576029,0.00012784185,0.000066199755,0.00024169021,0.00038390464,0.75474995,0.010351789,0.07000922,0.001360782,0.16042523],"study_design_scores_gemma":[0.000013828284,0.000028225953,0.00023472834,0.000004398022,0.000011664501,0.00004793445,0.000021989983,0.9870618,0.0010527271,0.009165256,0.0023507325,0.0000066713283],"about_ca_topic_score_codex":0.003035462,"about_ca_topic_score_gemma":0.004295152,"teacher_disagreement_score":0.003035462,"about_ca_system_score_codex":0.00074615836,"about_ca_system_score_gemma":0.0011709059,"threshold_uncertainty_score":0.008434832},"labels":[],"label_agreement":null},{"id":"W1995143551","doi":"10.1109/qsic.2013.38","title":"Interaction Models Matter in the Evaluation of Quality of Conceptual Models","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Completeness (order theory); Process (computing); Perspective (graphical); Conceptual model; Quality (philosophy); Quality assurance; Software engineering; Systems engineering; Data mining; Data science; Artificial intelligence; Engineering; Programming language; Database","score_opus":0.20139637090093987,"score_gpt":0.3930631360457694,"score_spread":0.19166676514482955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1995143551","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3789481,0.0022142185,0.5919378,0.0030588752,0.00013635324,0.00055826263,0.0005772802,0.0013782645,0.021190787],"genre_scores_gemma":[0.84662706,0.00041083017,0.15022925,0.00020886124,0.000036271387,0.00021126376,0.0005991998,0.00056131335,0.0011159601],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.88110584,0.053349443,0.008162289,0.005380131,0.050010737,0.0019915313],"domain_scores_gemma":[0.5173621,0.34153864,0.030084876,0.05830102,0.050312344,0.0024009955],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06037623,0.0011212736,0.0012527263,0.004702394,0.001376859,0.009693286,0.0018441018,0.0029424531,0.0027740544],"category_scores_gemma":[0.3037151,0.0011645682,0.0015683008,0.0029104634,0.003768035,0.013824004,0.004877403,0.002879765,0.00047508313],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030487727,0.00080073444,0.16860072,0.0032909464,0.0014283453,0.0012882174,0.02460708,0.10967374,0.030425578,0.24182768,0.004611232,0.410397],"study_design_scores_gemma":[0.00042007925,0.0038062853,0.099563755,0.0031081494,0.0019182293,0.0025868379,0.0122878,0.49358335,0.06686153,0.26678357,0.048399072,0.00068137],"about_ca_topic_score_codex":0.0034575337,"about_ca_topic_score_gemma":0.002870543,"teacher_disagreement_score":0.06037623,"about_ca_system_score_codex":0.0028203495,"about_ca_system_score_gemma":0.0030391726,"threshold_uncertainty_score":0.3193038},"labels":[],"label_agreement":null},{"id":"W1995440297","doi":"10.1145/1414004.1414049","title":"Enhancing predictive models using principal component analysis and search based metric selection","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Principal component analysis; Metric (unit); Computer science; Selection (genetic algorithm); Component (thermodynamics); Artificial intelligence; Data mining; Engineering","score_opus":0.04325257420471928,"score_gpt":0.28286649322276136,"score_spread":0.23961391901804208,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1995440297","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064800255,0.00048317504,0.9293247,0.0004624924,0.000047471905,0.00024418306,0.00021738773,0.0030337619,0.0013865355],"genre_scores_gemma":[0.55898553,0.00054554024,0.43723127,0.0001330902,0.00008746709,0.0005366329,0.0010754323,0.0003003069,0.0011047284],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99655294,0.0018705949,0.00019970906,0.00046128102,0.0007516301,0.00016387437],"domain_scores_gemma":[0.98347443,0.012595633,0.00076743925,0.000759916,0.0022377635,0.00016488985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006797767,0.002771593,0.0021339862,0.0055043856,0.0007331993,0.0022455938,0.0014159611,0.0010761489,0.0014847605],"category_scores_gemma":[0.027359793,0.0007722059,0.0015319414,0.004460712,0.00061790075,0.0024574776,0.0013495365,0.0016977127,0.0007594068],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000196163,0.00027276046,0.010739962,0.00016512249,0.00025377728,0.00009980765,0.00016433622,0.6965665,0.0017668819,0.003695667,0.001987459,0.28409144],"study_design_scores_gemma":[0.000009737522,0.00003787801,0.0006100752,0.000009510319,0.000025852345,0.000013342389,0.000012485732,0.99655956,0.00037383952,0.0021266644,0.00020868439,0.000012412421],"about_ca_topic_score_codex":0.014551927,"about_ca_topic_score_gemma":0.009200815,"teacher_disagreement_score":0.014551927,"about_ca_system_score_codex":0.0012349854,"about_ca_system_score_gemma":0.0021100491,"threshold_uncertainty_score":0.035950422},"labels":[],"label_agreement":null},{"id":"W1995584399","doi":"10.1142/s0218194009004489","title":"TEMPORAL SOFTWARE CHANGE PREDICTION USING NEURAL NETWORKS","year":2009,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Eclipse; Computer science; Software maintenance; Software; Software evolution; Dimension (graph theory); Plan (archaeology); Software development; Artificial neural network; Software engineering; Scale (ratio); Data science; Data mining; Artificial intelligence; Machine learning; Software construction; Operating system","score_opus":0.02304885103198641,"score_gpt":0.26661743351535744,"score_spread":0.24356858248337104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1995584399","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6086302,0.0013336147,0.3823237,0.0008181725,0.0001352968,0.000097341384,0.00086416083,0.0021449649,0.0036525433],"genre_scores_gemma":[0.97286224,0.00019150953,0.025467752,0.000050985967,0.000038460985,0.00004062344,0.000524827,0.00001910922,0.0008043729],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99937063,0.00015068227,0.000047661902,0.0002014372,0.00014053332,0.00008909003],"domain_scores_gemma":[0.99657613,0.002119984,0.0004707184,0.000150866,0.00060515606,0.00007710995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011453987,0.0006107556,0.00051251013,0.0017866505,0.00035185227,0.00078669825,0.00089767465,0.00085098966,0.0007826907],"category_scores_gemma":[0.006621259,0.00032888356,0.00053195615,0.0013860059,0.0003076756,0.0012054036,0.0004583937,0.00088240585,0.00017398153],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030387164,0.0001738258,0.022997694,0.00006651852,0.00010703318,0.00016050077,0.00007049583,0.82538176,0.0013282762,0.0009430996,0.0012601305,0.14720683],"study_design_scores_gemma":[0.0000017001076,0.0000065270183,0.00091641187,0.0000019290592,0.000004549194,0.0000055281994,0.000004047376,0.99841785,0.0001707988,0.0004176527,0.00005087367,0.0000021232606],"about_ca_topic_score_codex":0.027594378,"about_ca_topic_score_gemma":0.021903148,"teacher_disagreement_score":0.027594378,"about_ca_system_score_codex":0.0012973413,"about_ca_system_score_gemma":0.0005598939,"threshold_uncertainty_score":0.054867566},"labels":[],"label_agreement":null},{"id":"W1995653819","doi":"10.1109/msr.2013.6624021","title":"Understanding the evolution of Type-3 clones: An exploratory study","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Type (biology); Biology; Computer science; Programming language; Genetics; Gene","score_opus":0.10647273622953188,"score_gpt":0.29364954012610545,"score_spread":0.18717680389657357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1995653819","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9939956,0.00009045452,0.0052317698,0.000026464828,0.0000015027127,0.00006935818,0.0000881415,0.000033189554,0.00046356258],"genre_scores_gemma":[0.980492,0.00014068288,0.018212527,0.0000517386,0.0000055966593,0.000081064914,0.00037886683,0.00003952733,0.00059803517],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99716455,0.0010633097,0.00027356367,0.00056909665,0.00074058067,0.00018899734],"domain_scores_gemma":[0.95259804,0.030016527,0.0074690343,0.0044428557,0.0046453727,0.0008281844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035303393,0.000346929,0.0003929913,0.0018400431,0.0007644359,0.0010329344,0.00080698286,0.0008358864,0.0006408991],"category_scores_gemma":[0.021342957,0.00024749112,0.0005860586,0.0014687763,0.00072330807,0.0018331881,0.0010733089,0.0008127588,0.0001589898],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003936886,0.000769076,0.8369583,0.00030808634,0.00010266759,0.0022572496,0.02395338,0.0020582885,0.033320144,0.0012468172,0.00032816883,0.09830411],"study_design_scores_gemma":[0.000042124037,0.0015322367,0.9027314,0.00015845345,0.00022052205,0.0064963717,0.015840475,0.026126536,0.03650435,0.003440374,0.0067929174,0.00011431143],"about_ca_topic_score_codex":0.0013269483,"about_ca_topic_score_gemma":0.002559848,"teacher_disagreement_score":0.0035303393,"about_ca_system_score_codex":0.000622684,"about_ca_system_score_gemma":0.0005661341,"threshold_uncertainty_score":0.01867044},"labels":[],"label_agreement":null},{"id":"W1995760438","doi":"10.1109/saner.2015.7081861","title":"SPCP-Miner: A tool for mining code clones that are important for refactoring or tracking","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; Code (set theory); Computer science; Cloning (programming); Software; Software system; Software maintenance; Software engineering; Software evolution; Tracking (education); Tracking system; Programming language; Set (abstract data type); Software construction; Artificial intelligence","score_opus":0.17344531442390687,"score_gpt":0.3502397250376185,"score_spread":0.17679441061371165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1995760438","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.145144,0.0022840444,0.5940423,0.0008981166,0.00022548997,0.0020206799,0.032158233,0.21750881,0.0057183537],"genre_scores_gemma":[0.13597138,0.0005794726,0.81277865,0.00029268078,0.000046098903,0.0013183674,0.040237144,0.0052416706,0.0035345156],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99528164,0.0007542963,0.000718179,0.0011436027,0.0018807866,0.00022153166],"domain_scores_gemma":[0.97022635,0.0155019015,0.006069471,0.0037467584,0.0038286683,0.00062686374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043192483,0.0021729264,0.0014176732,0.012001366,0.0013032685,0.0020586164,0.0038862,0.0023392264,0.0027955167],"category_scores_gemma":[0.031012822,0.0014649036,0.0018583365,0.0064394567,0.0010391374,0.0037084965,0.002795065,0.0020825684,0.0023678918],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009175131,0.00104081,0.124917395,0.00495579,0.0008947631,0.0020464843,0.0030561455,0.023184836,0.030600293,0.0069633275,0.11425131,0.6871713],"study_design_scores_gemma":[0.00064344984,0.0010441014,0.057871286,0.0007916383,0.0006478115,0.0056455326,0.0013018618,0.67507774,0.08659575,0.020984082,0.14894532,0.00045143891],"about_ca_topic_score_codex":0.0042669033,"about_ca_topic_score_gemma":0.00929342,"teacher_disagreement_score":0.012001366,"about_ca_system_score_codex":0.00079444156,"about_ca_system_score_gemma":0.0033826174,"threshold_uncertainty_score":0.022842646},"labels":[],"label_agreement":null},{"id":"W1996180635","doi":"10.1145/1498926.1498927","title":"Deferring design pattern decisions and automating structural pattern changes using a design-pattern-based programming system","year":2009,"lang":"en","type":"article","venue":"ACM Transactions on Programming Languages and Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; University of Waterloo","funders":"University of Alberta","keywords":"Computer science; Software design pattern; Specification pattern; Design pattern; Generative Design; Coding (social sciences); Engineering design process; Structural pattern; Architecture; Context (archaeology); Software design; Software engineering; Programming language; Software development; Software","score_opus":0.05634059262171508,"score_gpt":0.3122899557454259,"score_spread":0.2559493631237108,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1996180635","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062205074,0.000024956173,0.9850932,0.0002238663,0.000025706262,0.00022298546,0.000041037263,0.0069382833,0.001209536],"genre_scores_gemma":[0.03730492,0.00006691869,0.9586576,0.00016958635,0.000014791747,0.00025463803,0.00017604283,0.0007819832,0.0025736028],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9942714,0.0020824408,0.0005824178,0.001204663,0.0015582994,0.00030085127],"domain_scores_gemma":[0.98583335,0.007460781,0.0012474293,0.0037189524,0.0014379006,0.00030159674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067922827,0.0015190353,0.0006233161,0.0011196737,0.00078039686,0.002565768,0.0025080505,0.0017469333,0.0024967943],"category_scores_gemma":[0.02123955,0.0011500239,0.0011459554,0.00080958346,0.0018864566,0.0033705235,0.002063744,0.0031397168,0.0013514111],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005613471,0.00070227135,0.0061732824,0.0006495865,0.00014073585,0.0012001804,0.0045869634,0.07407547,0.061901078,0.069414645,0.010053315,0.77054113],"study_design_scores_gemma":[0.00033545843,0.00063189096,0.0013481464,0.00030089432,0.00019982902,0.001472838,0.00049727695,0.720069,0.08847812,0.08953276,0.09692329,0.00021046496],"about_ca_topic_score_codex":0.001905701,"about_ca_topic_score_gemma":0.002368661,"teacher_disagreement_score":0.0067922827,"about_ca_system_score_codex":0.00094634364,"about_ca_system_score_gemma":0.0027391517,"threshold_uncertainty_score":0.035921395},"labels":[],"label_agreement":null},{"id":"W1996284583","doi":"10.1155/2013/952178","title":"A Granular Hierarchical Multiview Metrics Suite for Statecharts Quality","year":2013,"lang":"en","type":"article","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Suite; Process (computing); Quality (philosophy); Abstraction; Diagram; Quality assurance; Data mining; Software engineering; Programming language; Database","score_opus":0.018898756069967753,"score_gpt":0.30623282032282995,"score_spread":0.2873340642528622,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1996284583","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023891479,0.0001540457,0.97230136,0.000051121922,0.000013152129,0.00017346075,0.0007716817,0.0017979213,0.0008457636],"genre_scores_gemma":[0.2550669,0.00011169913,0.74189585,0.000025719557,0.000018235149,0.00037832782,0.0020137376,0.00017019955,0.00031937717],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9936127,0.0011265662,0.0008826545,0.0007688213,0.0033937923,0.0002155275],"domain_scores_gemma":[0.99140745,0.0025261906,0.0017023864,0.0012685825,0.0027511807,0.00034420678],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038867947,0.0013725001,0.0013943359,0.0102320975,0.0006344186,0.0026697754,0.0013514966,0.00091173116,0.0013061482],"category_scores_gemma":[0.018069472,0.00041756147,0.0016819654,0.003936794,0.0006221668,0.002706004,0.0018747975,0.0010593811,0.00040781312],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032419962,0.0003458655,0.04084819,0.0010271246,0.00041800467,0.00036660268,0.0010915573,0.24892479,0.03996293,0.048078842,0.002811053,0.6158008],"study_design_scores_gemma":[0.000023612394,0.00034891116,0.012381712,0.00013347385,0.0001042596,0.0002418737,0.00020922505,0.94263595,0.015408459,0.024752906,0.0036713434,0.00008820821],"about_ca_topic_score_codex":0.004051443,"about_ca_topic_score_gemma":0.0035382123,"teacher_disagreement_score":0.0102320975,"about_ca_system_score_codex":0.0014173493,"about_ca_system_score_gemma":0.0012026179,"threshold_uncertainty_score":0.020555556},"labels":[],"label_agreement":null},{"id":"W1996342119","doi":"10.1007/s10664-011-9188-2","title":"Introduction to the special issue on software repository mining in 2009","year":2011,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software engineering; Software; Data science; Operating system","score_opus":0.022772899001632055,"score_gpt":0.257169182466531,"score_spread":0.23439628346489894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1996342119","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012208449,0.053963155,0.051368836,0.09568327,0.75634664,0.0003228231,0.0035437685,0.001826418,0.035724275],"genre_scores_gemma":[0.0062692664,0.051010184,0.029010285,0.038056057,0.5995164,0.0002373334,0.0078927595,0.0022750092,0.2657328],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961427,0.00062309636,0.00054807856,0.0008314811,0.0016337732,0.00022093348],"domain_scores_gemma":[0.97921705,0.0059064445,0.0011497589,0.0018744979,0.0094211735,0.0024309333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062371315,0.0016737395,0.0026281804,0.008090917,0.0016585254,0.0073762364,0.0019392071,0.0029470983,0.056955993],"category_scores_gemma":[0.019519234,0.0009044778,0.0015214063,0.0060731196,0.0012312909,0.008996458,0.003911788,0.0057556983,0.039003015],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026684367,0.00003309595,0.00021017806,0.00025041803,0.000015005678,0.000050467028,0.00002573114,0.00014931754,0.00042108112,0.0015336248,0.9122657,0.08501866],"study_design_scores_gemma":[0.000007385538,0.00003726235,0.0005962593,0.00020937469,0.000012852863,0.00020763105,0.000028770179,0.00044569751,0.0003333879,0.0026455228,0.99545175,0.000024084346],"about_ca_topic_score_codex":0.0019072471,"about_ca_topic_score_gemma":0.005799961,"teacher_disagreement_score":0.056955993,"about_ca_system_score_codex":0.0024077077,"about_ca_system_score_gemma":0.003403722,"threshold_uncertainty_score":0.1905368},"labels":[],"label_agreement":null},{"id":"W1996993752","doi":"10.1109/tse.2013.28","title":"Early Detection of Collaboration Conflicts and Risks","year":2013,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Quality (philosophy); Control (management); Source code; Open source; Code (set theory); Software engineering; Risk analysis (engineering); Computer security; Data science; Software; Programming language; Business; Set (abstract data type); Artificial intelligence","score_opus":0.01665302119978797,"score_gpt":0.24331634647688724,"score_spread":0.22666332527709926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1996993752","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60524666,0.0009346738,0.37965217,0.0014716141,0.000087109715,0.00054779893,0.0006290685,0.005038901,0.006392099],"genre_scores_gemma":[0.8503848,0.00013440424,0.14804243,0.0000944984,0.00002037982,0.00011306365,0.00036163323,0.00014006734,0.0007086674],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9887066,0.0029253461,0.0008774013,0.0013851223,0.0054152207,0.00069039554],"domain_scores_gemma":[0.8991385,0.06119677,0.01749227,0.010174309,0.010059381,0.0019387883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00699047,0.0009032242,0.0006570572,0.0034389307,0.0010568554,0.0026907965,0.0020730728,0.0015166346,0.0014781457],"category_scores_gemma":[0.07130098,0.0010637621,0.0005304592,0.0012663908,0.0009568184,0.004116203,0.0037478625,0.0020095627,0.0003721414],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008819402,0.00047280206,0.49819484,0.00084599026,0.00021306833,0.0022940608,0.01118478,0.02339481,0.046761695,0.028181057,0.008632466,0.3789425],"study_design_scores_gemma":[0.00019154076,0.0007787377,0.16768098,0.0005710607,0.00034040815,0.003612338,0.006201289,0.6541431,0.06222674,0.074392766,0.029502587,0.00035852744],"about_ca_topic_score_codex":0.0028356048,"about_ca_topic_score_gemma":0.0031733299,"teacher_disagreement_score":0.00699047,"about_ca_system_score_codex":0.0010617257,"about_ca_system_score_gemma":0.0032263356,"threshold_uncertainty_score":0.036969602},"labels":[],"label_agreement":null},{"id":"W1997588820","doi":"10.5555/2486788.2487012","title":"Normalizing source code vocabulary to support program comprehension and software quality","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Program comprehension; Source code; Identifier; Normalization (sociology); Vocabulary; Software quality; Software maintenance; Natural language processing; Information retrieval; Static program analysis; Software; Artificial intelligence; Software development; Programming language; Software system; Linguistics","score_opus":0.05035563521686575,"score_gpt":0.3253980243126402,"score_spread":0.27504238909577444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1997588820","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16855133,0.0016419394,0.7933109,0.0004730197,0.00010746009,0.0006815457,0.00075410603,0.03020254,0.0042771404],"genre_scores_gemma":[0.4747864,0.00066061533,0.5150175,0.00025585733,0.000072411654,0.00063156313,0.0031573956,0.0037620952,0.0016562118],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99483854,0.0011985465,0.00075678364,0.0012585323,0.0017331964,0.0002144159],"domain_scores_gemma":[0.9752837,0.009687783,0.0039184405,0.005343705,0.0054373196,0.00032915032],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00334363,0.0009987982,0.0009398633,0.0031546883,0.0006365879,0.0020654928,0.0015498084,0.00071482634,0.0016170061],"category_scores_gemma":[0.041007243,0.00049187936,0.00080444646,0.0025474788,0.0009676131,0.0051239287,0.0026743577,0.0013663514,0.0009827932],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040581147,0.00038073552,0.0131941205,0.001300006,0.00011064186,0.00022045561,0.0036138927,0.01102462,0.17385018,0.008666522,0.0051202164,0.7821129],"study_design_scores_gemma":[0.00022222233,0.00085575285,0.03701678,0.0006545227,0.0006155427,0.0014656341,0.002820156,0.3080306,0.5279115,0.038194302,0.08185083,0.00036220934],"about_ca_topic_score_codex":0.003022281,"about_ca_topic_score_gemma":0.0029791528,"teacher_disagreement_score":0.00334363,"about_ca_system_score_codex":0.0010764283,"about_ca_system_score_gemma":0.002382082,"threshold_uncertainty_score":0.01768303},"labels":[],"label_agreement":null},{"id":"W1997670024","doi":"10.1155/2013/276105","title":"Extension of Object-Oriented Metrics Suite for Software Maintenance","year":2013,"lang":"en","type":"article","venue":"ISRN Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Maintainability; Computer science; Software metric; Suite; Software maintenance; Set (abstract data type); Measure (data warehouse); Data mining; Class (philosophy); Programming complexity; Software; Object-oriented programming; Software system; Software engineering; Software construction; Artificial intelligence; Programming language","score_opus":0.013296487586614638,"score_gpt":0.23849691223883435,"score_spread":0.2252004246522197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1997670024","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021018218,0.0011745925,0.946654,0.00040445046,0.00020664421,0.0018158914,0.004641693,0.015803043,0.008281494],"genre_scores_gemma":[0.105475046,0.0006944277,0.88064426,0.00010975209,0.00007761739,0.0018855594,0.0076298118,0.0011163758,0.0023671866],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98905295,0.0033178935,0.0019691454,0.00055011763,0.004882402,0.00022755882],"domain_scores_gemma":[0.98502004,0.00400106,0.0022236486,0.002363396,0.0060112793,0.00038054315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008457019,0.0017073528,0.0011122449,0.007361148,0.00071627845,0.0027397715,0.0015579477,0.0008534432,0.0022360682],"category_scores_gemma":[0.029102664,0.00046997645,0.0012041276,0.007768027,0.0003472286,0.0025336142,0.001800716,0.0014145485,0.0011822247],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003100257,0.00046372716,0.02643342,0.0017832922,0.0004652235,0.0005434596,0.0009558812,0.06053911,0.011764752,0.060923073,0.027317403,0.80850065],"study_design_scores_gemma":[0.00024113328,0.0017811654,0.04171119,0.0009457429,0.0005033362,0.0021017562,0.00038126035,0.5125048,0.02339698,0.11032416,0.30565915,0.00044932106],"about_ca_topic_score_codex":0.0035837865,"about_ca_topic_score_gemma":0.0033769766,"teacher_disagreement_score":0.008457019,"about_ca_system_score_codex":0.0010540929,"about_ca_system_score_gemma":0.0023141385,"threshold_uncertainty_score":0.044725537},"labels":[],"label_agreement":null},{"id":"W1998060612","doi":"10.1145/1920778.1920786","title":"Critic-proofing","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Usability; Heuristic evaluation; Computer science; Heuristic; Categorization; Software; Value (mathematics); Human–computer interaction; Artificial intelligence; Machine learning; Programming language","score_opus":0.012687353180009553,"score_gpt":0.2718507677242913,"score_spread":0.2591634145442817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1998060612","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009444076,0.00016965826,0.9473701,0.0019561239,0.00033960352,0.00096262293,0.00030977037,0.0049663284,0.03448169],"genre_scores_gemma":[0.2399137,0.00033037702,0.72536784,0.0011205883,0.00024658773,0.0011827669,0.00066252175,0.0018096368,0.029365974],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9819417,0.008338065,0.0010901423,0.0025370077,0.005202122,0.0008909964],"domain_scores_gemma":[0.8817171,0.07406679,0.005190747,0.022141546,0.015490043,0.0013938361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018093046,0.0023206428,0.0017053612,0.00399007,0.002830918,0.0046978816,0.004113573,0.0029204378,0.031901624],"category_scores_gemma":[0.13586423,0.0010636268,0.002288537,0.001789,0.005562796,0.006428599,0.0061379233,0.004596252,0.0054940404],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044805274,0.0003725045,0.005796929,0.0014842218,0.00022896366,0.0008263633,0.003362707,0.036572836,0.006761728,0.48039195,0.047390003,0.41636375],"study_design_scores_gemma":[0.0003295552,0.0004173229,0.0014883911,0.0008873738,0.00026381994,0.0010743778,0.0014144724,0.3924866,0.023181843,0.43123698,0.14700513,0.00021406458],"about_ca_topic_score_codex":0.002972991,"about_ca_topic_score_gemma":0.005655879,"teacher_disagreement_score":0.031901624,"about_ca_system_score_codex":0.0021963376,"about_ca_system_score_gemma":0.005057833,"threshold_uncertainty_score":0.10672158},"labels":[],"label_agreement":null},{"id":"W1998099342","doi":"10.1016/j.scico.2009.12.009","title":"Defining the meaning of tabular mathematical expressions","year":2010,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Expression (computer science); Schema (genetic algorithms); Meaning (existential); Set (abstract data type); Theoretical computer science; Semantics (computer science); Regular expression; Programming language; Artificial intelligence; Natural language processing; Information retrieval","score_opus":0.013313281107449868,"score_gpt":0.28131381459550436,"score_spread":0.26800053348805447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1998099342","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010470673,0.0023193357,0.9165828,0.0046045086,0.0016240232,0.000105630796,0.00079820934,0.0012737042,0.06222113],"genre_scores_gemma":[0.3127822,0.0031597756,0.6601753,0.0031191902,0.00147356,0.0004947604,0.0014910657,0.0017255046,0.0155786835],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9941806,0.0028564252,0.00077731407,0.00077207445,0.0009161204,0.00049736665],"domain_scores_gemma":[0.99239737,0.0034735568,0.00066418975,0.0013275373,0.001900945,0.0002364306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006309564,0.0011172311,0.001052199,0.00312063,0.0031693457,0.011188905,0.0029850774,0.002562885,0.010782554],"category_scores_gemma":[0.018625854,0.0013463252,0.0015820641,0.0040614954,0.009997412,0.021768847,0.004373636,0.005240186,0.0042638103],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012244642,0.000003575935,0.00004591506,0.00004256679,0.0000034331038,0.000025529906,0.00042149684,0.00018083838,0.00023639141,0.99257267,0.0014235433,0.005031715],"study_design_scores_gemma":[0.0000125128645,0.000012675994,0.000047833822,0.000105435276,0.000013193122,0.000101455014,0.00029631256,0.002544771,0.0008690405,0.9458282,0.05014727,0.000021302278],"about_ca_topic_score_codex":0.0013598398,"about_ca_topic_score_gemma":0.0010092191,"teacher_disagreement_score":0.011188905,"about_ca_system_score_codex":0.002339684,"about_ca_system_score_gemma":0.0017440435,"threshold_uncertainty_score":0.03607118},"labels":[],"label_agreement":null},{"id":"W1998265754","doi":"10.1109/ms.2006.105","title":"How are Java software developers using the Eclipse IDE?","year":2006,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":407,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Eclipse; Java; Computer science; Popularity; Software development; Development environment; Software engineering; Software; World Wide Web; Operating system","score_opus":0.02894594475689117,"score_gpt":0.25349211816357187,"score_spread":0.2245461734066807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1998265754","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8840393,0.007483756,0.009534806,0.03329183,0.00056755246,0.00004977906,0.00049935613,0.0005126783,0.06402092],"genre_scores_gemma":[0.97924954,0.0033410192,0.0027748453,0.002784298,0.00012934748,0.000034202065,0.00032714786,0.0001583081,0.011201251],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99234104,0.0030512223,0.0004597368,0.0010780976,0.0021712412,0.00089869194],"domain_scores_gemma":[0.97230023,0.010869801,0.0068133534,0.0018311609,0.0054850834,0.002700365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063014403,0.00025657946,0.00038422248,0.0014108112,0.0011516738,0.0051337336,0.00065046,0.0015226136,0.0034765617],"category_scores_gemma":[0.042271145,0.0004020494,0.00032478897,0.0013458034,0.0013147452,0.0080812825,0.0009257363,0.0013525602,0.0030566496],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017588721,0.00028031942,0.5961824,0.00018134053,0.00010954433,0.0011940065,0.026105072,0.00019074375,0.0024825074,0.0050130263,0.021473356,0.3466118],"study_design_scores_gemma":[0.000040512903,0.00023174843,0.6936818,0.00057483086,0.00016834142,0.008339372,0.10278235,0.0022539853,0.003144678,0.016415993,0.17214315,0.00022320042],"about_ca_topic_score_codex":0.0040589557,"about_ca_topic_score_gemma":0.0088157,"teacher_disagreement_score":0.0063014403,"about_ca_system_score_codex":0.00064498134,"about_ca_system_score_gemma":0.0010634465,"threshold_uncertainty_score":0.033325613},"labels":[],"label_agreement":null},{"id":"W1998569777","doi":"10.1109/icsme.2014.55","title":"Recommending Clones for Refactoring Using Design, Context, and History","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code refactoring; Computer science; clone (Java method); Cloning (programming); Context (archaeology); Software engineering; Software maintenance; Source code; Programming language; Code (set theory); Software development; Software","score_opus":0.1236162076083696,"score_gpt":0.30430109805012046,"score_spread":0.18068489044175085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1998569777","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6792535,0.0033674422,0.2990671,0.0007315899,0.00013985587,0.00037569593,0.0036960104,0.010434501,0.0029342738],"genre_scores_gemma":[0.73771966,0.00064677047,0.25115037,0.00016323205,0.00004939748,0.00017883086,0.007976691,0.0004018563,0.0017131167],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99804485,0.00033587005,0.00018783855,0.00067314645,0.00062902475,0.00012921004],"domain_scores_gemma":[0.98784953,0.0057736863,0.001499155,0.0012553647,0.0031548506,0.00046734067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002170528,0.0014296263,0.0009929693,0.0059423195,0.00063796353,0.0013840457,0.001211228,0.0014167036,0.0006640865],"category_scores_gemma":[0.014866036,0.00049222953,0.0011021128,0.002439682,0.00032757368,0.0019496917,0.0006802891,0.00092581793,0.0007664794],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033909347,0.00047303893,0.31036887,0.00059187383,0.00025728444,0.00057280855,0.00079523324,0.031756464,0.021095598,0.0008150254,0.011202839,0.62173194],"study_design_scores_gemma":[0.000080405945,0.00051491737,0.08712965,0.00022906398,0.0005024965,0.0010478051,0.00056045636,0.8679717,0.02393043,0.0033697095,0.014554064,0.0001092691],"about_ca_topic_score_codex":0.009541457,"about_ca_topic_score_gemma":0.021295037,"teacher_disagreement_score":0.009541457,"about_ca_system_score_codex":0.00084264233,"about_ca_system_score_gemma":0.0014031222,"threshold_uncertainty_score":0.0189718},"labels":[],"label_agreement":null},{"id":"W1998690866","doi":"10.1049/iet-sen:20070062","title":"Ontological approach for the semantic recovery of traceability links between software artefacts","year":2008,"lang":"en","type":"article","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Traceability; Computer science; Software engineering; Requirements traceability; Documentation; Ontology; Software development; Source code; Software construction; Software; Programming language","score_opus":0.060071436611381715,"score_gpt":0.2797311478187576,"score_spread":0.21965971120737587,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1998690866","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00260824,0.00020924666,0.99131525,0.00085242424,0.00009837244,0.00008849363,0.00012041871,0.00027378058,0.0044336785],"genre_scores_gemma":[0.081713185,0.0007890573,0.9110514,0.00039716798,0.000109649,0.00033697984,0.00068671413,0.00012255326,0.004793256],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9940619,0.001827065,0.00077381916,0.0006027191,0.002408401,0.00032608112],"domain_scores_gemma":[0.99081236,0.0032325045,0.0007601893,0.0035142757,0.0014352861,0.00024537573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058116745,0.001013955,0.00080839515,0.0063725286,0.0027653147,0.0052965614,0.0037407225,0.003379422,0.0021948884],"category_scores_gemma":[0.0149161145,0.00089811767,0.004037884,0.004674609,0.0050944914,0.010571711,0.006948262,0.0059159384,0.0007558841],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004447357,0.00012946862,0.000767602,0.00027034065,0.00012814174,0.0005186903,0.002435724,0.009025099,0.0029770983,0.91300523,0.002294034,0.068404],"study_design_scores_gemma":[0.000031717667,0.00004732855,0.0007099608,0.00032700459,0.0002390551,0.0006189693,0.0012073369,0.113165475,0.006389057,0.7844093,0.09275212,0.00010261873],"about_ca_topic_score_codex":0.008008258,"about_ca_topic_score_gemma":0.008397128,"teacher_disagreement_score":0.008008258,"about_ca_system_score_codex":0.003341992,"about_ca_system_score_gemma":0.0046448144,"threshold_uncertainty_score":0.030735433},"labels":[],"label_agreement":null},{"id":"W1999284448","doi":"10.1016/j.scico.2012.01.004","title":"Taupe : Visualizing and analyzing eye-tracking data","year":2012,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Program comprehension; Documentation; Software engineering; Plug-in; Extensibility; Eye tracking; Software; Process (computing); Human–computer interaction; Usability; Context (archaeology); Software system; Data science; Programming language; Artificial intelligence","score_opus":0.04910834378222158,"score_gpt":0.3630185884685466,"score_spread":0.31391024468632506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999284448","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039565347,0.0011349573,0.770238,0.00028153742,0.00033964784,0.0007965616,0.019833822,0.16115059,0.0066595087],"genre_scores_gemma":[0.14465137,0.0014164721,0.82036227,0.000270034,0.000115365016,0.001644385,0.010188664,0.009313615,0.012037803],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996518,0.00004852123,0.0000330845,0.00009713111,0.00012657657,0.000042820488],"domain_scores_gemma":[0.99892163,0.00054885825,0.00011109828,0.0001199282,0.00021941226,0.000079083155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00081827544,0.0015394086,0.0008171012,0.004012585,0.0005032646,0.0019716385,0.0010111916,0.00090190594,0.022525625],"category_scores_gemma":[0.0031021023,0.00049118197,0.0007802032,0.0018109134,0.0002320419,0.0015676572,0.0014119598,0.0007498043,0.0044094184],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014546314,0.00023341752,0.0053658416,0.0022430103,0.00054009503,0.0009919942,0.0016175385,0.0056686476,0.17582819,0.0042939857,0.11286549,0.68889713],"study_design_scores_gemma":[0.0006219826,0.0008041883,0.04496331,0.0006307804,0.00070426584,0.0035026367,0.002203327,0.43911302,0.30476722,0.028332073,0.17381041,0.0005468339],"about_ca_topic_score_codex":0.0037601818,"about_ca_topic_score_gemma":0.0075950976,"teacher_disagreement_score":0.022525625,"about_ca_system_score_codex":0.0003106336,"about_ca_system_score_gemma":0.00080228614,"threshold_uncertainty_score":0.07535577},"labels":[],"label_agreement":null},{"id":"W1999457095","doi":"10.1109/scam.2014.11","title":"Automatic Identification of Important Clones for Refactoring and Tracking","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; Computer science; Software evolution; clone (Java method); Software maintenance; Cloning (programming); Identification (biology); Software system; Software; Software engineering; Programming language; Tracking (education); Code (set theory); Tracking system; Artificial intelligence; Software construction; Biology; Genetics","score_opus":0.02363146154358002,"score_gpt":0.2901312652968421,"score_spread":0.26649980375326204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999457095","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6932459,0.0011266131,0.29256472,0.0002806212,0.00006891779,0.0006086694,0.0026484302,0.0073887534,0.002067378],"genre_scores_gemma":[0.602311,0.00045669076,0.38745534,0.00009208851,0.00003647376,0.000287861,0.0063384078,0.000589695,0.0024323843],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971469,0.00033552083,0.00037983037,0.0008783297,0.0011019113,0.00015748777],"domain_scores_gemma":[0.9769495,0.008437245,0.00579574,0.002919366,0.005435652,0.0004626102],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022713903,0.0006015099,0.0008043177,0.006643347,0.00073049014,0.0014940341,0.0011615155,0.00086384115,0.0009279175],"category_scores_gemma":[0.019167734,0.0004033006,0.00074211525,0.003634001,0.00035950006,0.0016729342,0.0009989158,0.0006568697,0.0005817212],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003447511,0.00027178042,0.26054075,0.0006457312,0.0001405493,0.0009952111,0.0021835433,0.005669944,0.05714144,0.0018904536,0.0034679065,0.666708],"study_design_scores_gemma":[0.0001153185,0.0008088693,0.40925965,0.00042712622,0.0006719977,0.003961387,0.0018778001,0.40483585,0.14246804,0.008269763,0.027076237,0.00022809672],"about_ca_topic_score_codex":0.0039280755,"about_ca_topic_score_gemma":0.00668895,"teacher_disagreement_score":0.006643347,"about_ca_system_score_codex":0.00047067474,"about_ca_system_score_gemma":0.0014231143,"threshold_uncertainty_score":0.0120123625},"labels":[],"label_agreement":null},{"id":"W1999875458","doi":"10.1145/1865841.1865856","title":"Semantic comparison of structured visual dataflow programs","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Dataflow; Computer science; Programming language; Debugger; Visual programming language; Dataflow architecture; Debugging","score_opus":0.022503892553128662,"score_gpt":0.33092810911672577,"score_spread":0.3084242165635971,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999875458","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068256594,0.00006805448,0.9242718,0.00013973584,0.00003151638,0.00013571396,0.00034729292,0.0024669783,0.0042822585],"genre_scores_gemma":[0.39457402,0.00010413707,0.60018235,0.00012276934,0.000022625456,0.00028250794,0.0016505148,0.00089061295,0.0021704226],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985536,0.0003534599,0.00013980737,0.00026991375,0.00056889834,0.00011446211],"domain_scores_gemma":[0.9964116,0.0016099218,0.00030230734,0.0005485476,0.0009946886,0.00013289282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017889055,0.00037399933,0.0004754521,0.0023229988,0.0005347209,0.0018081117,0.0011844581,0.0006231317,0.0032339308],"category_scores_gemma":[0.0084879715,0.00032812692,0.00093964394,0.0011763127,0.0012459329,0.0033095935,0.0017919759,0.00060474745,0.00030280274],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008381077,0.00027353226,0.005591566,0.00062289526,0.00009573732,0.00042584728,0.0020822438,0.048287988,0.045581885,0.5519066,0.004336606,0.3399569],"study_design_scores_gemma":[0.00012972149,0.00027354507,0.002097732,0.00015693158,0.00008818331,0.00030541574,0.00091047445,0.35503393,0.11342327,0.49282303,0.034670863,0.00008690521],"about_ca_topic_score_codex":0.0010358676,"about_ca_topic_score_gemma":0.0010741702,"teacher_disagreement_score":0.0032339308,"about_ca_system_score_codex":0.0011203032,"about_ca_system_score_gemma":0.0011183934,"threshold_uncertainty_score":0.010818601},"labels":[],"label_agreement":null},{"id":"W1999910737","doi":"10.1109/msr.2013.6624026","title":"A contextual approach towards more accurate duplicate bug report detection","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":129,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Data deduplication; Software bug; BitTorrent tracker; Security bug; Software regression; Vocabulary; Software; Context (archaeology); Android (operating system); Debugging; Data science; World Wide Web; Software engineering; Software quality; Software development; Artificial intelligence; Database; Computer security; Software security assurance; Information security; Eye tracking","score_opus":0.028980335204376376,"score_gpt":0.277534679319666,"score_spread":0.24855434411528965,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999910737","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07060857,0.0059234616,0.9025314,0.0012252711,0.00038364364,0.0004653271,0.0016322643,0.013021759,0.0042083464],"genre_scores_gemma":[0.2942924,0.0012523693,0.69842327,0.0004081794,0.0005403596,0.00031759831,0.0023551325,0.0006356477,0.0017750205],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993145,0.0020175113,0.00082586124,0.002085721,0.0015873482,0.00033843218],"domain_scores_gemma":[0.9709781,0.009129726,0.003530783,0.009070032,0.006727538,0.00056385604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056424784,0.0014358766,0.0024474382,0.010461057,0.0016073617,0.0027948364,0.0019463782,0.0014757611,0.0019879236],"category_scores_gemma":[0.033508733,0.0008369978,0.0011638087,0.0069083776,0.0010112061,0.004384126,0.0039064297,0.001661125,0.0013344759],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006559823,0.00054841925,0.029533906,0.0014970317,0.0002732697,0.00075488666,0.002765972,0.013302628,0.0488185,0.009191772,0.01667841,0.8759793],"study_design_scores_gemma":[0.00032715927,0.0013242239,0.054764424,0.000766136,0.001383554,0.004553251,0.0038493313,0.6219202,0.12487808,0.047430154,0.13812597,0.000677459],"about_ca_topic_score_codex":0.0045753713,"about_ca_topic_score_gemma":0.010101948,"teacher_disagreement_score":0.010461057,"about_ca_system_score_codex":0.0008452418,"about_ca_system_score_gemma":0.0023409887,"threshold_uncertainty_score":0.029840589},"labels":[],"label_agreement":null},{"id":"W2000540309","doi":"10.1145/1512762.1512771","title":"Do software libraries evolve differently than applications?","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software; Software evolution; Software engineering; Software development; Programming language; Software construction","score_opus":0.01868000737844582,"score_gpt":0.264405528907546,"score_spread":0.24572552152910015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2000540309","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9812841,0.0017686988,0.0048169913,0.004290728,0.000052293526,0.000023076496,0.00025677623,0.00014372391,0.007363582],"genre_scores_gemma":[0.99622715,0.0005854458,0.0013939908,0.0005593485,0.00005363074,0.000014548549,0.0002746511,0.00004684386,0.00084422866],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99638283,0.0014435268,0.00022487731,0.0007775403,0.0009164251,0.00025486606],"domain_scores_gemma":[0.9398708,0.029484099,0.019626355,0.0040156827,0.0055324803,0.0014706011],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029923494,0.0003415644,0.00049466995,0.0026111559,0.0005878001,0.0032893098,0.0006383485,0.0017028099,0.0015886887],"category_scores_gemma":[0.052466203,0.00030038887,0.00028334715,0.0056756153,0.0021081094,0.008711075,0.0006625067,0.00080179056,0.00072581315],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000099609875,0.00011738547,0.87447554,0.00014236162,0.00009841578,0.00016382048,0.0020503306,0.0013271138,0.0019384706,0.0050527947,0.001170293,0.11336388],"study_design_scores_gemma":[0.000009113465,0.00008657607,0.9769667,0.000060900973,0.00004647085,0.00030539188,0.0016822492,0.0050390144,0.0016065411,0.009517747,0.0046432316,0.000036087797],"about_ca_topic_score_codex":0.0048534726,"about_ca_topic_score_gemma":0.004965356,"teacher_disagreement_score":0.0048534726,"about_ca_system_score_codex":0.0016166731,"about_ca_system_score_gemma":0.0006204176,"threshold_uncertainty_score":0.015825212},"labels":[],"label_agreement":null},{"id":"W2000558260","doi":"10.1504/ijlsm.2011.038601","title":"Distinguishing the indistinguishable: exploring differences in supply chain software packages using centering resonance text analysis","year":2011,"lang":"en","type":"article","venue":"International Journal of Logistics Systems and Management","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Documentation; Computer science; Software; Software package; Software engineering; Selection (genetic algorithm); Data science; Artificial intelligence; Programming language","score_opus":0.11486444187254485,"score_gpt":0.28535020753828544,"score_spread":0.17048576566574059,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2000558260","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94192946,0.00032638933,0.05176485,0.00046789309,0.00001902255,0.00011917553,0.00016857844,0.00014244697,0.00506221],"genre_scores_gemma":[0.94056886,0.00017671101,0.058102045,0.00006107244,0.00001543433,0.000089595414,0.00039884358,0.00007152126,0.00051593623],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99520254,0.0019476061,0.00052852655,0.00054988527,0.0016115822,0.00015983447],"domain_scores_gemma":[0.9595387,0.030540457,0.004355173,0.0013427056,0.0038092735,0.0004136609],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043756966,0.00038526708,0.00032758433,0.008726146,0.0011051818,0.0030105382,0.0007239439,0.0006544936,0.0014092388],"category_scores_gemma":[0.036330994,0.00022423179,0.00039747008,0.005620937,0.0022544768,0.006359161,0.0016852964,0.0009231212,0.00034756723],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081503816,0.00025339428,0.094823785,0.0012408325,0.000111575646,0.001136397,0.18239988,0.0016769684,0.040421434,0.030675057,0.0016552071,0.6447905],"study_design_scores_gemma":[0.0001680954,0.0012917368,0.39100453,0.0018070597,0.0004524561,0.0035252257,0.2815918,0.059211224,0.040042512,0.17226823,0.048132632,0.00050450966],"about_ca_topic_score_codex":0.0013890689,"about_ca_topic_score_gemma":0.0017638634,"teacher_disagreement_score":0.008726146,"about_ca_system_score_codex":0.0009210368,"about_ca_system_score_gemma":0.0011065054,"threshold_uncertainty_score":0.023141146},"labels":[],"label_agreement":null},{"id":"W2000624622","doi":"10.1016/j.jss.2013.10.044","title":"Visualizing protected variations in evolving software designs","year":2013,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Fonds de Recherche du Québec - Santé; École de technologie supérieure","keywords":"Computer science; Software; Software visualization; Task (project management); Software design; Software engineering; Software analytics; Software construction; Software evolution; Visualization; Software development; Human–computer interaction; Data science; Data mining; Systems engineering; Engineering; Programming language","score_opus":0.02924881058216502,"score_gpt":0.2726796775483301,"score_spread":0.2434308669661651,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2000624622","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6643793,0.0006707669,0.31678346,0.00043511254,0.00013342522,0.000083024985,0.0005882905,0.009089157,0.00783738],"genre_scores_gemma":[0.85318667,0.0001973792,0.14278638,0.000038148348,0.000022308124,0.000039942155,0.00050167897,0.0011727747,0.002054773],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994554,0.00012413891,0.000033022592,0.00008174339,0.00025208626,0.000053669086],"domain_scores_gemma":[0.9950257,0.0025661616,0.0006284429,0.00074063183,0.00072156783,0.00031755207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000922653,0.000662217,0.0003643196,0.0025548027,0.0005288926,0.0014214172,0.0008086868,0.0009793922,0.0038184172],"category_scores_gemma":[0.0057605905,0.00053809764,0.0005071295,0.001534505,0.0007100659,0.0016944661,0.001331949,0.0012609708,0.0003345837],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017025525,0.00044014197,0.038437765,0.00093450863,0.00018205098,0.0027556098,0.011060442,0.2121181,0.13226414,0.06234258,0.010252588,0.5275095],"study_design_scores_gemma":[0.000098424534,0.0004539628,0.018433148,0.0001601994,0.00014322132,0.0013661747,0.0012182011,0.8800627,0.033937834,0.041772313,0.02222959,0.00012424668],"about_ca_topic_score_codex":0.0023152693,"about_ca_topic_score_gemma":0.003577004,"teacher_disagreement_score":0.0038184172,"about_ca_system_score_codex":0.00051510753,"about_ca_system_score_gemma":0.0006650619,"threshold_uncertainty_score":0.012773871},"labels":[],"label_agreement":null},{"id":"W2001176736","doi":"10.4018/jitwe.2007070101","title":"Modeling Defects In E-Projects","year":2007,"lang":"en","type":"article","venue":"International Journal of Information Technology and Web Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software engineering; Software; Software development; Context (archaeology); Field (mathematics); Software development process; Process (computing); Web application; Systems engineering; Data science; World Wide Web; Engineering","score_opus":0.006940783084816767,"score_gpt":0.2411291371439762,"score_spread":0.23418835405915944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2001176736","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6135839,0.00032349155,0.37639078,0.00056034053,0.000035487978,0.00020527694,0.0004360517,0.00048466845,0.007979988],"genre_scores_gemma":[0.9690488,0.00029273255,0.025851706,0.0000403521,0.0000148295685,0.00021749016,0.00025381192,0.000053983844,0.004226358],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987047,0.00045843638,0.000100636076,0.000262185,0.00025134403,0.00022278553],"domain_scores_gemma":[0.9913737,0.0054861982,0.0016339306,0.0004656697,0.00076417514,0.00027631494],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002376329,0.00086096924,0.00066828117,0.0018300463,0.00039673288,0.0017474323,0.0021196988,0.0024471076,0.0019001946],"category_scores_gemma":[0.011721489,0.0005997553,0.0010300673,0.0013433932,0.0011843403,0.002388536,0.001337038,0.0009929636,0.0003314865],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041057316,0.00014340747,0.009232996,0.000046223216,0.000027306076,0.00014994184,0.00018175592,0.9502153,0.0006081739,0.03086618,0.0002851952,0.00820242],"study_design_scores_gemma":[0.000011345555,0.000043405606,0.0012280751,0.000007082392,0.000011603802,0.00003268544,0.000049945636,0.98913795,0.00016123941,0.008991536,0.00031751813,0.0000076095525],"about_ca_topic_score_codex":0.011365931,"about_ca_topic_score_gemma":0.005084628,"teacher_disagreement_score":0.011365931,"about_ca_system_score_codex":0.0014425672,"about_ca_system_score_gemma":0.0010838113,"threshold_uncertainty_score":0.022599518},"labels":[],"label_agreement":null},{"id":"W2001292754","doi":"10.1145/1028664.1028670","title":"JQuery","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Programming language; Context (archaeology); Code (set theory); Class (philosophy); Source code; Field (mathematics); Tree (set theory); World Wide Web; Artificial intelligence","score_opus":0.014656477622904767,"score_gpt":0.2542762898968811,"score_spread":0.23961981227397636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2001292754","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005601795,0.0019580557,0.25135654,0.0030755152,0.0012883008,0.0008418806,0.07456701,0.54174393,0.119566984],"genre_scores_gemma":[0.05515359,0.0030571038,0.21128315,0.006662773,0.000748359,0.0015688057,0.23499814,0.32237303,0.16415511],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99692255,0.000404049,0.00034282837,0.00063597906,0.0013497116,0.00034480702],"domain_scores_gemma":[0.99694854,0.0012960237,0.0002063557,0.0006221107,0.0007072649,0.00021976467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033337541,0.0017791307,0.0012685964,0.0024194482,0.0015860739,0.0053375796,0.004732475,0.0032654868,0.11937103],"category_scores_gemma":[0.008877361,0.0018172459,0.0019858105,0.002115652,0.0011566068,0.0077799396,0.007039695,0.0031221374,0.082071565],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070235675,0.000116162664,0.0018023596,0.0014872595,0.00013567666,0.00050336856,0.0011060251,0.0008097854,0.013936165,0.03321758,0.856254,0.0899293],"study_design_scores_gemma":[0.000118525095,0.000054190736,0.0010509765,0.00014816542,0.000035634657,0.0004874209,0.00026579204,0.0038538359,0.006055122,0.015180033,0.972644,0.00010638847],"about_ca_topic_score_codex":0.006668975,"about_ca_topic_score_gemma":0.00908593,"teacher_disagreement_score":0.11937103,"about_ca_system_score_codex":0.0012852271,"about_ca_system_score_gemma":0.0017846293,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2001599472","doi":"10.1007/s11219-010-9128-1","title":"An industrial case study of classifier ensembles for locating software defects","year":2011,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Sabancı Üniversitesi","keywords":"Computer science; Software; Classifier (UML); False alarm; Software quality; Data mining; Domain (mathematical analysis); Ensemble learning; Embedded software; Machine learning; Artificial intelligence; Reliability engineering; Software development; Engineering; Operating system","score_opus":0.2236840213549557,"score_gpt":0.3779922822970453,"score_spread":0.1543082609420896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2001599472","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92701125,0.0003644177,0.0681788,0.0003270322,0.00004784302,0.00018857248,0.0001990621,0.00054325425,0.003139879],"genre_scores_gemma":[0.9472494,0.000100067016,0.05086465,0.000035583893,0.00001198632,0.000035296976,0.00015688197,0.000045595472,0.0015004908],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9974268,0.0009018117,0.00016063856,0.00039168616,0.00092753547,0.00019146343],"domain_scores_gemma":[0.98154956,0.011882583,0.0006998042,0.0025172767,0.0029522656,0.00039849692],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038478682,0.0006649975,0.0006426304,0.0015652886,0.0010413994,0.0011298647,0.0019106513,0.0020956686,0.0015853302],"category_scores_gemma":[0.013429443,0.00031128168,0.0005537263,0.0015837582,0.0006263921,0.0010696582,0.0007609113,0.00087570254,0.00050237385],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023078131,0.0038892815,0.111221366,0.00083137787,0.0003215936,0.009728532,0.002760554,0.21614277,0.0465234,0.004251115,0.0075352346,0.59448695],"study_design_scores_gemma":[0.0003210168,0.0040647783,0.033370852,0.00010182652,0.00029792258,0.006298017,0.0022498749,0.87591225,0.06423763,0.0044330982,0.008612101,0.00010066007],"about_ca_topic_score_codex":0.003948894,"about_ca_topic_score_gemma":0.006650547,"teacher_disagreement_score":0.003948894,"about_ca_system_score_codex":0.0007788204,"about_ca_system_score_gemma":0.00059896696,"threshold_uncertainty_score":0.020349681},"labels":[],"label_agreement":null},{"id":"W2001935537","doi":"10.5555/2486788.2486997","title":"V:ISSUE:LIZER: exploring requirements clarification in online communication over time","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Domain (mathematical analysis); Requirements analysis; Software engineering; World Wide Web; Software","score_opus":0.0800850046072287,"score_gpt":0.31178199114644706,"score_spread":0.23169698653921836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2001935537","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19181764,0.0012611342,0.38520703,0.009138942,0.0016015222,0.0029027658,0.028036244,0.15400034,0.22603437],"genre_scores_gemma":[0.37360427,0.0009964292,0.4257882,0.0019946964,0.00041501503,0.0022606212,0.024295986,0.019305123,0.15133967],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.9993711,0.00026877792,0.000027028302,0.000074898904,0.00019495263,0.00006332899],"domain_scores_gemma":[0.994997,0.0036594607,0.0001710089,0.00037718794,0.00049080147,0.00030450328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016301607,0.00065167865,0.00039891733,0.0014467157,0.0007579994,0.0018042825,0.0009659031,0.0013991619,0.07136194],"category_scores_gemma":[0.008243026,0.00027179264,0.00053987314,0.0007526962,0.00036214027,0.0041625267,0.0024304087,0.00087859854,0.0121384235],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007800145,0.00047630063,0.00518007,0.0015464072,0.00005941276,0.0014562717,0.023024905,0.0016990481,0.026114488,0.00877188,0.5454436,0.3854476],"study_design_scores_gemma":[0.00022704486,0.0007602197,0.022374973,0.00093646673,0.00007175851,0.0016954999,0.014897217,0.03610002,0.021354811,0.013217916,0.88806474,0.00029926337],"about_ca_topic_score_codex":0.0014904198,"about_ca_topic_score_gemma":0.003834301,"teacher_disagreement_score":0.07136194,"about_ca_system_score_codex":0.00047084395,"about_ca_system_score_gemma":0.00042044974,"threshold_uncertainty_score":0.23872942},"labels":[],"label_agreement":null},{"id":"W2002413229","doi":"10.1016/j.datak.2012.09.005","title":"Comparing functionality of software systems: An ontological approach","year":2012,"lang":"en","type":"article","venue":"Data & Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software system; Software engineering; Adaptation (eye); Function (biology); USable; Software development; Software; Software metric; Software quality; Point (geometry); Software construction; Function point; Quality (philosophy); Software sizing; World Wide Web; Programming language","score_opus":0.13962881508491679,"score_gpt":0.31436609711026914,"score_spread":0.17473728202535235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002413229","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16646494,0.0005929579,0.8076092,0.0012604959,0.00007145385,0.00031424165,0.0009699265,0.0006111232,0.022105671],"genre_scores_gemma":[0.61139673,0.0005047121,0.3848816,0.00016661505,0.00003227603,0.00018695554,0.0016415454,0.00013221275,0.0010573148],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.992191,0.0030510798,0.0011860228,0.00064001424,0.0024701706,0.00046179522],"domain_scores_gemma":[0.98269147,0.0105126295,0.0010270475,0.0035612213,0.0017499741,0.0004577295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007802832,0.0006332292,0.0006532074,0.00963484,0.0019320605,0.0070888223,0.002571203,0.001677986,0.00187896],"category_scores_gemma":[0.025253268,0.0005982196,0.0026994017,0.005578896,0.0046209213,0.011985433,0.0034268573,0.0013716918,0.00029469965],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016627609,0.00023236826,0.018761946,0.0006292909,0.00026612228,0.0005868811,0.007017516,0.014035983,0.012522172,0.7852103,0.0011157149,0.15945555],"study_design_scores_gemma":[0.000054570064,0.0003088634,0.017943965,0.00084850774,0.0009353033,0.00092948833,0.011061015,0.09883818,0.012604464,0.8239443,0.032386944,0.00014442587],"about_ca_topic_score_codex":0.0062020007,"about_ca_topic_score_gemma":0.0045564296,"teacher_disagreement_score":0.00963484,"about_ca_system_score_codex":0.0020271393,"about_ca_system_score_gemma":0.0030405226,"threshold_uncertainty_score":0.041265786},"labels":[],"label_agreement":null},{"id":"W2002449827","doi":"10.1109/wcre.2012.54","title":"The Secret Life of Patches: A Firefox Case Study","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Casual; Process (computing); Quality (philosophy); Computer science; Process management; Code review; Open source; Business; Software quality; Software development; Software; Political science","score_opus":0.03498157899868887,"score_gpt":0.2979043337525235,"score_spread":0.26292275475383464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002449827","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98450893,0.00037198546,0.0055324016,0.0021431523,0.00002857029,0.00016546961,0.00006836254,0.00008667751,0.007094431],"genre_scores_gemma":[0.98589003,0.0004143936,0.0084267575,0.00033371823,0.000038916372,0.00010000563,0.00009266861,0.00008790137,0.004615618],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.993171,0.004100682,0.0002827392,0.00048943795,0.0013371797,0.00061900145],"domain_scores_gemma":[0.9538669,0.034189392,0.0035364635,0.0027044527,0.002976902,0.0027258447],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010832357,0.0003791472,0.0003451451,0.0016874415,0.005775554,0.0025287515,0.0016028387,0.0024468203,0.0015855792],"category_scores_gemma":[0.027608035,0.00044935054,0.00041952552,0.0014438459,0.0039396808,0.0037341728,0.00289011,0.001865997,0.00039558555],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007696425,0.002440092,0.13149638,0.0007824048,0.00007877112,0.067787044,0.5783881,0.004759294,0.013448639,0.023143228,0.0128108375,0.16409554],"study_design_scores_gemma":[0.000383476,0.0033207315,0.18167293,0.0012933573,0.00021396112,0.054674614,0.47062457,0.02371515,0.020258462,0.017742226,0.2257839,0.00031661463],"about_ca_topic_score_codex":0.0084901955,"about_ca_topic_score_gemma":0.017708525,"teacher_disagreement_score":0.010832357,"about_ca_system_score_codex":0.003465906,"about_ca_system_score_gemma":0.0022069425,"threshold_uncertainty_score":0.057287633},"labels":[],"label_agreement":null},{"id":"W2002559933","doi":"","title":"Shuffling and randomization for scalable source code clone detection","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Concordia University","funders":"","keywords":"Shuffling; Computer science; Scalability; Code (set theory); Source code; clone (Java method); Machine learning; Theoretical computer science; Programming language; Database","score_opus":0.012515661390295772,"score_gpt":0.2410320729598843,"score_spread":0.22851641156958852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002559933","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11561419,0.00086953,0.83761173,0.0004044903,0.0001631691,0.00042582778,0.0009003516,0.04230721,0.0017034535],"genre_scores_gemma":[0.4909476,0.00018147373,0.5039414,0.00028968224,0.00012981822,0.00047695896,0.0018881093,0.0010144875,0.0011306001],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99288046,0.0023279123,0.0006116593,0.001687836,0.0021440827,0.00034792302],"domain_scores_gemma":[0.9707865,0.011397879,0.0029290742,0.012111221,0.0022686212,0.0005067255],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050967247,0.0009873708,0.00096084207,0.0034539504,0.00078931596,0.0017141274,0.0022060962,0.0011384283,0.0016730808],"category_scores_gemma":[0.02509513,0.00047559405,0.0009750173,0.0028893268,0.0011529168,0.0034933446,0.0022706338,0.0015260704,0.0013454049],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012546385,0.0006428459,0.034103822,0.0004476044,0.00034948142,0.0005231708,0.00045370197,0.03833896,0.11557002,0.013072673,0.019358369,0.7758848],"study_design_scores_gemma":[0.00019426882,0.00073608087,0.00862013,0.000073450625,0.000115416275,0.00091401325,0.00020063046,0.78288376,0.16780974,0.02285893,0.015463696,0.00012995522],"about_ca_topic_score_codex":0.0010036281,"about_ca_topic_score_gemma":0.00128512,"teacher_disagreement_score":0.0050967247,"about_ca_system_score_codex":0.00076725514,"about_ca_system_score_gemma":0.0015978507,"threshold_uncertainty_score":0.026954353},"labels":[],"label_agreement":null},{"id":"W2002641269","doi":"10.1007/s10664-014-9315-y","title":"An empirical study on the importance of source code entities for requirements traceability","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Traceability; Source code; Weighting; Search engine indexing; Rank (graph theory); Data mining; Software engineering; Programming language","score_opus":0.05999686482834988,"score_gpt":0.3489345172867725,"score_spread":0.2889376524584226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002641269","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9968644,0.00007780567,0.0013146653,0.00011237478,0.000003660893,0.000030697072,0.00003662547,0.000008162946,0.0015515782],"genre_scores_gemma":[0.9986204,0.00004920134,0.0010145019,0.000020140626,0.0000036487654,0.000013528911,0.00006786043,0.000006771797,0.00020391987],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99054915,0.0052949777,0.00071197166,0.0007286245,0.0023461604,0.00036912464],"domain_scores_gemma":[0.3697808,0.5703203,0.029988302,0.011228925,0.015800359,0.0028813074],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013134783,0.00031313722,0.00017528802,0.0030472584,0.00080040656,0.0014995472,0.0009260275,0.0008791077,0.002844926],"category_scores_gemma":[0.21636586,0.00027645135,0.00032293488,0.0030462793,0.0015832394,0.004169205,0.0011633161,0.001961907,0.00024936374],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007239827,0.0029902891,0.9129549,0.0003344015,0.00009686486,0.00034721536,0.009467962,0.0012336025,0.0036939266,0.0036830816,0.00037379737,0.06409977],"study_design_scores_gemma":[0.00007548805,0.0012600981,0.96458125,0.00017065475,0.00017558095,0.00062030944,0.013458465,0.010476842,0.004544277,0.0021485623,0.002455292,0.000033258948],"about_ca_topic_score_codex":0.0033265215,"about_ca_topic_score_gemma":0.0049050865,"teacher_disagreement_score":0.013134783,"about_ca_system_score_codex":0.0011256079,"about_ca_system_score_gemma":0.001974818,"threshold_uncertainty_score":0.06946415},"labels":[],"label_agreement":null},{"id":"W2002706478","doi":"10.1109/vlhcc.2012.6344511","title":"Usable results from the field of API usability: A systematic mapping and further analysis","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"USable; Computer science; Usability; Field (mathematics); Software engineering; Data science; Software; Mainstream; Application programming interface; Human–computer interaction; World Wide Web; Programming language","score_opus":0.025553484003445304,"score_gpt":0.2727688466331918,"score_spread":0.24721536262974647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002706478","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42919034,0.44645676,0.05252734,0.005921343,0.0007548251,0.018872134,0.011404575,0.0003675422,0.034505192],"genre_scores_gemma":[0.7319085,0.17423666,0.066504896,0.001768494,0.00023826087,0.017955346,0.00509108,0.00024526092,0.0020516103],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.90184873,0.04337591,0.02642918,0.0031763953,0.023297455,0.0018722726],"domain_scores_gemma":[0.49773225,0.354481,0.030683136,0.016863015,0.0985553,0.0016852593],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08025326,0.0016385644,0.003847806,0.0835616,0.0018221741,0.0062285163,0.0015958925,0.0013843635,0.0038429874],"category_scores_gemma":[0.2929508,0.0009137013,0.0042577134,0.046267956,0.0021672712,0.0077960533,0.0044571036,0.0011632944,0.00071259163],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009808668,0.0005314278,0.057507295,0.2747609,0.0057614227,0.0014376781,0.04704514,0.00083534804,0.00438027,0.0073606446,0.0056364243,0.5937625],"study_design_scores_gemma":[0.00083467126,0.0040374775,0.29320034,0.3764646,0.033477265,0.0034546987,0.12506093,0.0021764215,0.016009167,0.022219349,0.12251742,0.0005476759],"about_ca_topic_score_codex":0.0020079634,"about_ca_topic_score_gemma":0.0033326042,"teacher_disagreement_score":0.91974676,"about_ca_system_score_codex":0.003805268,"about_ca_system_score_gemma":0.012847984,"threshold_uncertainty_score":0.4244249},"labels":[],"label_agreement":null},{"id":"W200284635","doi":"","title":"Incremental Effort Prediction Models in Agile Development using Radial Basis Functions.","year":2007,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Extreme programming; Agile software development; Computer science; Extreme programming practices; Software engineering; Software development; Scratch; Overhead (engineering); Software; Agile Unified Process; Estimation; Agile usability engineering; Software development process; Industrial engineering; Systems engineering; Engineering; Programming language","score_opus":0.017740558077816375,"score_gpt":0.23429649239955824,"score_spread":0.21655593432174186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W200284635","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17688856,0.0005160482,0.82014215,0.00031163075,0.0000339531,0.000094486386,0.000102983104,0.0005124016,0.0013979187],"genre_scores_gemma":[0.89383763,0.00029330954,0.10362106,0.00004741357,0.000020962656,0.00018431865,0.00021033395,0.00006254731,0.0017224913],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981186,0.0010889764,0.00007668115,0.00023239526,0.00034592202,0.00013752429],"domain_scores_gemma":[0.98974633,0.0075919623,0.00087213813,0.00043270263,0.0011786083,0.00017823536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054053874,0.0010580674,0.0009065921,0.0013159654,0.00032514246,0.0011680068,0.0018391573,0.0011642388,0.00083705486],"category_scores_gemma":[0.018702772,0.0005264942,0.0006911844,0.0011721501,0.0005161788,0.0017550798,0.0008658123,0.0013223825,0.00048787895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025061797,0.00016329437,0.0061152675,0.00006302534,0.00008204265,0.000080665624,0.00019296755,0.91424364,0.0005554379,0.0053555323,0.0005650495,0.072332494],"study_design_scores_gemma":[0.000004053032,0.000031187323,0.00060036447,0.000005592512,0.000008055406,0.0000101827345,0.000011639978,0.99765766,0.00011198933,0.0014661211,0.000086100445,0.0000070979377],"about_ca_topic_score_codex":0.0105528785,"about_ca_topic_score_gemma":0.006462179,"teacher_disagreement_score":0.0105528785,"about_ca_system_score_codex":0.0010313699,"about_ca_system_score_gemma":0.00081231125,"threshold_uncertainty_score":0.028586745},"labels":[],"label_agreement":null},{"id":"W2002850079","doi":"10.5555/2486788.2486895","title":"Reverb: recommending code-related web pages","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"World Wide Web; Computer science; Documentation; Web page; Code (set theory); Source code; Web development; Field (mathematics); Programming language; Set (abstract data type)","score_opus":0.02251973246789852,"score_gpt":0.25598035986049666,"score_spread":0.23346062739259815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002850079","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7299609,0.004909009,0.17374983,0.001570345,0.0003560895,0.0008793515,0.003194635,0.0724312,0.012948623],"genre_scores_gemma":[0.62550277,0.0009186252,0.35641226,0.0003229679,0.000103883474,0.00018040994,0.0037639411,0.0008217871,0.011973447],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989334,0.00039292645,0.000054366606,0.00021442532,0.00035082377,0.000054083324],"domain_scores_gemma":[0.99008256,0.0057818666,0.00082814624,0.0013446801,0.0014240247,0.0005386872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001900628,0.0010701157,0.0006860423,0.002871119,0.0005078904,0.0012100443,0.0012104362,0.0012745326,0.004858607],"category_scores_gemma":[0.014480724,0.0006188461,0.0003587156,0.0012266162,0.00024693675,0.0014940248,0.00063241436,0.0009059207,0.0033165687],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013046109,0.0022054533,0.054104228,0.0013541344,0.00023021361,0.00083361863,0.0009854906,0.026242742,0.03329915,0.0011401853,0.05551067,0.8227895],"study_design_scores_gemma":[0.000519698,0.0023429685,0.049635768,0.000319715,0.00030837665,0.0015289254,0.00084560463,0.82629645,0.04829756,0.0018984816,0.06774897,0.00025754067],"about_ca_topic_score_codex":0.009402058,"about_ca_topic_score_gemma":0.019430505,"teacher_disagreement_score":0.009402058,"about_ca_system_score_codex":0.00043544933,"about_ca_system_score_gemma":0.0008899892,"threshold_uncertainty_score":0.018694699},"labels":[],"label_agreement":null},{"id":"W2002926403","doi":"10.1145/2557833.2557856","title":"1st international workshop on conducting empirical studies in industry (CESI 2013)","year":2014,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Empirical research; Theme (computing); Engineering management; Software; Computer science; Engineering; Software quality; Software engineering; Software development; World Wide Web; Mathematics","score_opus":0.14868082622399667,"score_gpt":0.37487886635105605,"score_spread":0.22619804012705938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002926403","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0067113354,0.054585464,0.42994395,0.21716538,0.09825449,0.004743426,0.007947195,0.011051386,0.16959742],"genre_scores_gemma":[0.04801716,0.059978053,0.5015258,0.05767318,0.029004557,0.010572781,0.032742813,0.01335337,0.24713232],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.93909776,0.034066733,0.0037732741,0.0047752773,0.014296107,0.0039908523],"domain_scores_gemma":[0.8765588,0.039737266,0.0030399875,0.018420873,0.043621354,0.01862184],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.11592721,0.0029476953,0.003030715,0.0068437555,0.0034668914,0.016598027,0.006891645,0.008569893,0.07590887],"category_scores_gemma":[0.09409026,0.0021251875,0.0036825251,0.00520898,0.0044982582,0.013931873,0.024078695,0.013368282,0.04080151],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000263772,0.0003985797,0.0007957272,0.001003101,0.00007271125,0.0002792807,0.0025523296,0.00076431374,0.0023753918,0.025152434,0.7067791,0.25956324],"study_design_scores_gemma":[0.00004609175,0.00009174625,0.0012335064,0.0016897522,0.000024494744,0.00021741379,0.0012281543,0.000676052,0.0008539284,0.015302736,0.97857136,0.00006478306],"about_ca_topic_score_codex":0.004612041,"about_ca_topic_score_gemma":0.0057357014,"teacher_disagreement_score":0.8840728,"about_ca_system_score_codex":0.005643974,"about_ca_system_score_gemma":0.019634258,"threshold_uncertainty_score":0.613089},"labels":[],"label_agreement":null},{"id":"W2002988936","doi":"10.1109/empire.2012.6347683","title":"LASR: A tool for large scale annotation of software requirements","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Annotation; Computer science; Set (abstract data type); Software; Scale (ratio); Adjudication; Information retrieval; Field (mathematics); Data science; Software engineering; Artificial intelligence","score_opus":0.030287812572943048,"score_gpt":0.3096044863192998,"score_spread":0.27931667374635677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002988936","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0038570727,0.00013147657,0.821585,0.00035705583,0.00010140294,0.0007891665,0.0060521835,0.16076101,0.0063656396],"genre_scores_gemma":[0.03392016,0.00024233767,0.91748166,0.0003755512,0.000074479685,0.001997829,0.024631921,0.0127343265,0.008541725],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9833407,0.00713974,0.002179859,0.0017523781,0.0051623667,0.00042491147],"domain_scores_gemma":[0.94290584,0.03310187,0.004093224,0.0118696755,0.0071409224,0.0008884217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017365018,0.0026972562,0.0014370566,0.009792806,0.001992998,0.003638233,0.0037941134,0.0025070128,0.023441909],"category_scores_gemma":[0.045065165,0.0018834763,0.0018465135,0.0044218823,0.0016033158,0.0065338034,0.006992095,0.0034347253,0.018386364],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063240295,0.00051747164,0.0025437397,0.004091276,0.00020780522,0.001756091,0.008105901,0.008107118,0.060131546,0.029172353,0.2623919,0.6223424],"study_design_scores_gemma":[0.0003725798,0.00044258664,0.007412288,0.001156252,0.00013717312,0.0022964838,0.0029409064,0.13781571,0.0647506,0.04813906,0.73391503,0.0006213191],"about_ca_topic_score_codex":0.0029460646,"about_ca_topic_score_gemma":0.0057393303,"teacher_disagreement_score":0.023441909,"about_ca_system_score_codex":0.0015446139,"about_ca_system_score_gemma":0.0037005537,"threshold_uncertainty_score":0.091836095},"labels":[],"label_agreement":null},{"id":"W2003040216","doi":"10.4236/jsea.2012.57051","title":"Improving Class Cohesion Measurement: Towards a Novel Approach Using Hierarchical Clustering","year":2012,"lang":"en","type":"article","venue":"Journal of Software Engineering and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Innovation and Economic Development Trois Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Cohesion (chemistry); Computer science; Cluster analysis; Object-oriented programming; Data mining; Artificial intelligence; Programming language","score_opus":0.038973867852384506,"score_gpt":0.25715591585108916,"score_spread":0.21818204799870466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003040216","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01371426,0.00036824867,0.98306006,0.00011142078,0.000038822804,0.00016590378,0.00016377404,0.0014787766,0.0008986915],"genre_scores_gemma":[0.13053007,0.00024152215,0.8667661,0.00005708009,0.00006214798,0.00025795418,0.0006747509,0.00043204305,0.0009783943],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99424607,0.0013456118,0.00039685686,0.001264936,0.0024955883,0.0002509197],"domain_scores_gemma":[0.9895726,0.003207239,0.0014201823,0.0016193249,0.0039327242,0.00024790998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031094651,0.001524382,0.0017137611,0.008989454,0.0011266652,0.0022137468,0.0027763192,0.0014377396,0.0012050977],"category_scores_gemma":[0.013341646,0.0006597533,0.001293888,0.0062962295,0.0006913768,0.0032315091,0.002272136,0.0013111283,0.0010004552],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013782397,0.00033771255,0.01184015,0.0007085986,0.00044693568,0.00012279813,0.0013066526,0.04242263,0.030873887,0.008562102,0.006653368,0.8965873],"study_design_scores_gemma":[0.000083797466,0.00030557098,0.026247686,0.0001794457,0.0005353408,0.00039472748,0.0009712216,0.8914073,0.031877372,0.026540207,0.021235941,0.00022141603],"about_ca_topic_score_codex":0.007820735,"about_ca_topic_score_gemma":0.008257884,"teacher_disagreement_score":0.008989454,"about_ca_system_score_codex":0.0012615632,"about_ca_system_score_gemma":0.0017606171,"threshold_uncertainty_score":0.016444623},"labels":[],"label_agreement":null},{"id":"W2003042352","doi":"10.1504/ijasm.2015.068605","title":"How project duration, upfront costs and uncertainty interact and impact on software development productivity? A simulation approach","year":2015,"lang":"en","type":"article","venue":"International Journal of Agile Systems and Management","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia","funders":"","keywords":"Productivity; Software development; Project management; Duration (music); Software project management; Software; Computer science; Team software process; Risk analysis (engineering); Process management; Engineering; Systems engineering; Software development process; Business; Economics; Software construction","score_opus":0.03190719140086141,"score_gpt":0.30368363046906355,"score_spread":0.27177643906820215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003042352","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9789073,0.00023561362,0.012297025,0.00040258162,0.000010739361,0.00003814008,0.00014862219,0.000022245107,0.007937845],"genre_scores_gemma":[0.99737144,0.0001522108,0.0017489018,0.000020653295,0.0000033651038,0.000037652655,0.000043835626,0.0000063216767,0.0006156771],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978321,0.0015886682,0.000058599737,0.00012016492,0.000158029,0.00024236694],"domain_scores_gemma":[0.9730584,0.023658974,0.0016044219,0.0003772032,0.0007248568,0.00057613483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030777948,0.00046643193,0.0004396545,0.00083436567,0.0003085689,0.0020518398,0.00058859476,0.0010207769,0.0031581966],"category_scores_gemma":[0.017806808,0.00035900399,0.0007383853,0.0009809041,0.0005268179,0.0021059848,0.0009065562,0.0007589021,0.00021033084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007166093,0.00053608837,0.09807848,0.00017459052,0.00023261526,0.00030124,0.0010208659,0.8548784,0.001972034,0.021623347,0.000573464,0.019892316],"study_design_scores_gemma":[0.00008248503,0.00046845584,0.02684726,0.00005416242,0.00016032162,0.000089107074,0.0008507203,0.9585988,0.0009193459,0.010621305,0.0012540612,0.000054036744],"about_ca_topic_score_codex":0.0070623476,"about_ca_topic_score_gemma":0.0051722145,"teacher_disagreement_score":0.0070623476,"about_ca_system_score_codex":0.0014067856,"about_ca_system_score_gemma":0.001710471,"threshold_uncertainty_score":0.016277075},"labels":[],"label_agreement":null},{"id":"W2003234666","doi":"10.1109/ccece.2014.6900911","title":"Cooperative based software clustering on dependency graphs","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"University of Waterloo","keywords":"Cluster analysis; Computer science; Dependency (UML); Software; Data mining; Graph; Dependency graph; Software system; Theoretical computer science; Artificial intelligence; Machine learning; Programming language","score_opus":0.016277362558926717,"score_gpt":0.2548321198509188,"score_spread":0.23855475729199208,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003234666","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02273359,0.00017537067,0.974857,0.000116575095,0.000011170388,0.00007717549,0.00009858948,0.00055781705,0.0013728105],"genre_scores_gemma":[0.4967987,0.00044307194,0.49712443,0.000101690144,0.000050743758,0.00027514546,0.0010141064,0.00035072878,0.0038413657],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978067,0.0009186822,0.000085729014,0.00056232995,0.00048051547,0.00014595741],"domain_scores_gemma":[0.9937402,0.0034969426,0.0004939926,0.001044937,0.0010226408,0.00020129426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002105529,0.0010896127,0.00094081607,0.0038584778,0.001425642,0.0013582504,0.002110248,0.001190419,0.0015943256],"category_scores_gemma":[0.009767324,0.00077816384,0.0014414281,0.0030535124,0.0012177496,0.002967035,0.0022418364,0.0009110725,0.00059750705],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000110968125,0.00010092708,0.002698563,0.00019100061,0.00012810047,0.00016792271,0.0004783445,0.77256507,0.003626714,0.044386424,0.0033282537,0.17221767],"study_design_scores_gemma":[0.0000074417376,0.000025729449,0.0006017142,0.000010395557,0.00002342177,0.000048309474,0.000067695815,0.9639243,0.0014192846,0.03247529,0.0013845708,0.000011865761],"about_ca_topic_score_codex":0.009322344,"about_ca_topic_score_gemma":0.01134799,"teacher_disagreement_score":0.009322344,"about_ca_system_score_codex":0.0019474678,"about_ca_system_score_gemma":0.0014150786,"threshold_uncertainty_score":0.01853615},"labels":[],"label_agreement":null},{"id":"W2003468239","doi":"10.1016/s0164-1212(02)00055-9","title":"API documentation with executable examples","year":2003,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Documentation; Executable; Internal documentation; Computer science; Software documentation; Software engineering; Test (biology); Consistency (knowledge bases); Technical documentation; Test case; Programming language; Software; Software system; Software development; Software development process; Artificial intelligence; Software construction","score_opus":0.015681775056133943,"score_gpt":0.2449624474642674,"score_spread":0.22928067240813346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003468239","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018215256,0.00028541996,0.762783,0.0010496114,0.00041378962,0.00029133068,0.0038990283,0.1448878,0.068174824],"genre_scores_gemma":[0.18866776,0.0007501189,0.6683721,0.00047918805,0.00020787335,0.0005313373,0.016506627,0.050462827,0.07402212],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975623,0.0005416815,0.0004034301,0.00022633874,0.0011236498,0.00014264719],"domain_scores_gemma":[0.988853,0.0047086235,0.00056990737,0.0033944135,0.0023181953,0.00015572958],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016605646,0.0012852754,0.000688527,0.0022331306,0.0010031583,0.0026539303,0.0015834703,0.0018328216,0.058579378],"category_scores_gemma":[0.019946696,0.0015550265,0.0008448585,0.0018165004,0.00055732916,0.0043353206,0.0025561298,0.00291017,0.023792984],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005432281,0.00046805528,0.0024127339,0.0014529818,0.00006128873,0.0022649697,0.00184548,0.009317282,0.016317276,0.07581813,0.26040405,0.6290945],"study_design_scores_gemma":[0.00025273024,0.00016210144,0.002251495,0.0012616548,0.00013026086,0.005229765,0.000385347,0.1350576,0.094221815,0.080343515,0.6804875,0.00021619535],"about_ca_topic_score_codex":0.00099164,"about_ca_topic_score_gemma":0.0016551409,"teacher_disagreement_score":0.058579378,"about_ca_system_score_codex":0.00041549138,"about_ca_system_score_gemma":0.0010639197,"threshold_uncertainty_score":0.19596756},"labels":[],"label_agreement":null},{"id":"W2003484266","doi":"10.1109/icsm.2010.5609696","title":"2D and 3D visualizations in WikiDev2.0","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Visualization; Computer science; Zoom; Interactivity; Human–computer interaction; Software; Software visualization; Information visualization; Data visualization; Process (computing); World Wide Web; Software development; Component-based software engineering; Engineering; Artificial intelligence","score_opus":0.012335786041077084,"score_gpt":0.29628553723393913,"score_spread":0.28394975119286203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003484266","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06716129,0.0014293473,0.85768414,0.0009529177,0.0004376752,0.00045558007,0.008734108,0.041121453,0.022023503],"genre_scores_gemma":[0.2340114,0.001320821,0.7444894,0.00026854972,0.0001267597,0.0008906033,0.0058502825,0.0063104974,0.0067316177],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985353,0.00052968244,0.00010630324,0.00017692157,0.0005628237,0.000089045454],"domain_scores_gemma":[0.9946603,0.003438229,0.00030295594,0.00080276915,0.00060574926,0.00018999095],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021015343,0.0011356296,0.0008154744,0.0041325213,0.0006973515,0.0037576745,0.0008511024,0.0010330318,0.008704025],"category_scores_gemma":[0.0099293515,0.0008684555,0.0007533754,0.0031580552,0.0006300155,0.0036935953,0.0037217103,0.0013594137,0.0016638776],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018707099,0.00034392346,0.011371414,0.0032123278,0.00023667401,0.0020866927,0.016823357,0.034299176,0.05412393,0.041590396,0.08138092,0.7526606],"study_design_scores_gemma":[0.00044771345,0.0005080906,0.030781927,0.0017365724,0.00019634818,0.0033500947,0.0052971584,0.17240705,0.0772304,0.069228336,0.63782024,0.0009961127],"about_ca_topic_score_codex":0.0021838956,"about_ca_topic_score_gemma":0.0030518044,"teacher_disagreement_score":0.008704025,"about_ca_system_score_codex":0.00027732516,"about_ca_system_score_gemma":0.00069146376,"threshold_uncertainty_score":0.029117882},"labels":[],"label_agreement":null},{"id":"W2003876621","doi":"10.1109/saner.2015.7081812","title":"Mining Multi-level API Usage Patterns","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Application programming interface; Computer science; Software; Usage data; World Wide Web; Operating system","score_opus":0.15566951263966797,"score_gpt":0.32547756277600765,"score_spread":0.16980805013633968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003876621","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.71613604,0.0013004122,0.26584578,0.00031721292,0.000063499756,0.0004034914,0.0074533634,0.0058789765,0.0026012077],"genre_scores_gemma":[0.82823515,0.00037854488,0.15817846,0.000111117115,0.000046340778,0.00044351319,0.010555954,0.0004450219,0.0016059651],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948408,0.0006893238,0.0008086427,0.0015220213,0.0017559662,0.0003831809],"domain_scores_gemma":[0.9865335,0.005151985,0.002837271,0.0022810746,0.0027517823,0.00044424852],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018423718,0.0012853777,0.0012029719,0.01071627,0.00075070164,0.0015225101,0.0017523684,0.0011472824,0.0007499165],"category_scores_gemma":[0.012857233,0.00071026204,0.0014853991,0.007297988,0.00043747952,0.0023051975,0.0018659746,0.00095789146,0.0006945408],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058456534,0.000819039,0.40193856,0.0013283353,0.00087245135,0.002694647,0.0020376306,0.015721967,0.04945382,0.002506116,0.0082138525,0.513829],"study_design_scores_gemma":[0.00008350194,0.0005459799,0.31151232,0.00027404728,0.0007334449,0.0069639687,0.0022003893,0.59201705,0.053427435,0.011821545,0.02022636,0.00019401568],"about_ca_topic_score_codex":0.003379467,"about_ca_topic_score_gemma":0.0053665787,"teacher_disagreement_score":0.01071627,"about_ca_system_score_codex":0.00042760946,"about_ca_system_score_gemma":0.0010134238,"threshold_uncertainty_score":0.009743452},"labels":[],"label_agreement":null},{"id":"W2003920582","doi":"10.1145/1414004.1414063","title":"Analysis of the reliability of a subset of change metrics for defect prediction","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Eclipse; Computer science; Reliability (semiconductor); Precision and recall; Software metric; Software quality; Data mining; Stability (learning theory); Software; Metric (unit); Software bug; Reliability engineering; Machine learning; Artificial intelligence; Software development; Engineering; Programming language","score_opus":0.08058338221403603,"score_gpt":0.28971770584256257,"score_spread":0.20913432362852652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003920582","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.934215,0.0010533779,0.061097004,0.00018948193,0.000059896123,0.00007012725,0.0011625321,0.001347729,0.00080486946],"genre_scores_gemma":[0.98881185,0.000117989366,0.009317461,0.000021078447,0.000032938657,0.000026986756,0.0014337987,0.00009374576,0.0001441382],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9931452,0.0028039645,0.0005810107,0.0010807771,0.0020982735,0.00029071947],"domain_scores_gemma":[0.86888015,0.09730582,0.0067261625,0.01204016,0.01373073,0.0013169401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0117278425,0.0017946272,0.001679805,0.0055153854,0.0004809652,0.0010856212,0.0008189986,0.0011396273,0.00052097416],"category_scores_gemma":[0.07028923,0.00043805037,0.0011991508,0.002592206,0.00042700756,0.001590651,0.0006507514,0.0012364167,0.0005249613],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002432037,0.0006051718,0.5991499,0.00041464408,0.0021771695,0.0004031459,0.00043001265,0.14614592,0.028770903,0.0003035667,0.002482236,0.21668537],"study_design_scores_gemma":[0.00003712189,0.0014705654,0.15641934,0.000040146795,0.00040458338,0.00047849838,0.0000944445,0.8259324,0.013808449,0.0006248729,0.00063174986,0.000057811878],"about_ca_topic_score_codex":0.003187781,"about_ca_topic_score_gemma":0.0028427935,"teacher_disagreement_score":0.0117278425,"about_ca_system_score_codex":0.00041512027,"about_ca_system_score_gemma":0.0005510225,"threshold_uncertainty_score":0.06202346},"labels":[],"label_agreement":null},{"id":"W2004258507","doi":"10.1109/scam.2012.26","title":"Improving Bug Location Using Binary Class Relationships","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Information retrieval; Source code; Class (philosophy); Code (set theory); Search engine indexing; Ranking (information retrieval); Rank (graph theory); Binary number; Empirical research; Data mining; Programming language; Artificial intelligence","score_opus":0.06652502757580461,"score_gpt":0.28970288696433244,"score_spread":0.22317785938852783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2004258507","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73817563,0.0032016097,0.19265684,0.0009785463,0.0001041902,0.00039592746,0.0049910555,0.04707265,0.012423587],"genre_scores_gemma":[0.8439672,0.00033976696,0.1483704,0.00012161762,0.000041551953,0.00011495917,0.0044720834,0.00073816365,0.0018341853],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958075,0.0009751568,0.0003341811,0.00091730413,0.0017435369,0.00022235246],"domain_scores_gemma":[0.97441417,0.013168572,0.005818217,0.003552747,0.0025869398,0.00045932224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041429983,0.0010708577,0.0010013655,0.009961523,0.0008174002,0.0021165563,0.0019597514,0.0010154707,0.0025034167],"category_scores_gemma":[0.030636687,0.00041683487,0.0009586671,0.006823583,0.00075224595,0.0062887603,0.0018456418,0.001034723,0.0012779039],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007199418,0.0008069473,0.10277688,0.0008631749,0.00012207426,0.00009390433,0.00089416554,0.015568607,0.0151861925,0.0045184493,0.012560265,0.84588945],"study_design_scores_gemma":[0.00034448312,0.0016876465,0.12363881,0.00020534666,0.0004266155,0.00074463297,0.0014016982,0.79253274,0.040065635,0.016227262,0.022409463,0.0003156982],"about_ca_topic_score_codex":0.012488494,"about_ca_topic_score_gemma":0.014463872,"teacher_disagreement_score":0.012488494,"about_ca_system_score_codex":0.0012252604,"about_ca_system_score_gemma":0.0018977869,"threshold_uncertainty_score":0.024831593},"labels":[],"label_agreement":null},{"id":"W2004295584","doi":"10.1007/s10115-008-0187-6","title":"The analysis and management of non-canonical requirement specifications through a belief integration game","year":2009,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Viewpoints; Set (abstract data type); Analogy; Outcome (game theory); Theoretical computer science; Mathematics; Programming language; Mathematical economics","score_opus":0.029788926324930975,"score_gpt":0.29060839348361633,"score_spread":0.26081946715868537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2004295584","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22430095,0.00007782927,0.76565534,0.0014215837,0.000021181051,0.0002615402,0.000073421834,0.00023454532,0.007953606],"genre_scores_gemma":[0.90457416,0.000049502978,0.093454234,0.000097444194,0.000014282265,0.00019855461,0.00006668956,0.00003682958,0.0015083352],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913573,0.0051434836,0.00029535836,0.0006475792,0.0016348871,0.0009214281],"domain_scores_gemma":[0.93206185,0.059495714,0.0025234134,0.0022929043,0.0023197345,0.0013063777],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011062417,0.00082390016,0.0011503828,0.0009806326,0.0010308061,0.003613183,0.0026991912,0.0022679316,0.0030634021],"category_scores_gemma":[0.052659664,0.0009803039,0.0012003495,0.0008931826,0.0028925582,0.0062580234,0.00298561,0.003069863,0.00024357563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094525353,0.0006934424,0.004187309,0.0002114531,0.00017279627,0.00061573135,0.0025677995,0.3663638,0.0054324814,0.57020986,0.001604236,0.046995807],"study_design_scores_gemma":[0.00007085908,0.000094876734,0.00028218297,0.000015569802,0.000030154568,0.000028352892,0.00017868895,0.8833995,0.00067769113,0.114797145,0.0004030083,0.000021914391],"about_ca_topic_score_codex":0.0072438554,"about_ca_topic_score_gemma":0.006178808,"teacher_disagreement_score":0.011062417,"about_ca_system_score_codex":0.0025221389,"about_ca_system_score_gemma":0.0037616647,"threshold_uncertainty_score":0.058504343},"labels":[],"label_agreement":null},{"id":"W2004586836","doi":"10.1145/2347696.2347698","title":"Technical debt in software development","year":2012,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Software engineering; Debt; Software; Engineering management; Computer science; Software development; Software quality; Engineering; Systems engineering; Business; Finance","score_opus":0.021891034526968448,"score_gpt":0.2594726961300906,"score_spread":0.23758166160312216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2004586836","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17086531,0.03277584,0.36808696,0.080363266,0.0014092381,0.00023910972,0.00024965208,0.00085110404,0.3451595],"genre_scores_gemma":[0.9378669,0.006457031,0.03187903,0.0023220473,0.00042591273,0.00022610404,0.000095729614,0.0003291853,0.020398017],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98875976,0.006988925,0.00076924707,0.0007170759,0.0022430883,0.0005219814],"domain_scores_gemma":[0.97211546,0.017329507,0.0033626338,0.0027065878,0.0030036797,0.00148211],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009961194,0.0005448366,0.0004941902,0.0024933051,0.004240759,0.008931414,0.0011295349,0.0029984259,0.0038876995],"category_scores_gemma":[0.04265603,0.0007314953,0.00036270433,0.004775977,0.010606241,0.018604608,0.008805452,0.0050049243,0.0006828542],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000056273413,0.00003746262,0.0029742087,0.00024619332,0.000010282963,0.00044127807,0.015557073,0.0018410721,0.00046979223,0.9066515,0.007927432,0.06378739],"study_design_scores_gemma":[0.000032043197,0.000048448193,0.002136529,0.00042799034,0.000011913869,0.0007194393,0.0057331305,0.004798783,0.00049531244,0.8604369,0.12512395,0.000035558747],"about_ca_topic_score_codex":0.0036307406,"about_ca_topic_score_gemma":0.0025964996,"teacher_disagreement_score":0.009961194,"about_ca_system_score_codex":0.006259839,"about_ca_system_score_gemma":0.0028094237,"threshold_uncertainty_score":0.052680433},"labels":[],"label_agreement":null},{"id":"W2004681822","doi":"10.1109/ictai.2006.106","title":"Software Maintenance: Similarity and Inclusion of Rules in Knowledge Extraction","year":2006,"lang":"en","type":"article","venue":"Proceedings - International Conference on Tools with Artificial Intelligence, TAI","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; National Research Council Institute for Biodiagnostics; University of Alberta","funders":"","keywords":"Maintainability; Software maintenance; Computer science; Software engineering; Verification and validation; Software development; Software metric; Software; Software construction; Set (abstract data type); Data mining; Software sizing; Similarity (geometry); Software system; Artificial intelligence; Engineering; Programming language","score_opus":0.05730947791003136,"score_gpt":0.321796407484553,"score_spread":0.26448692957452163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2004681822","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31981927,0.002363569,0.6643601,0.00057092245,0.000070875234,0.0011807381,0.001422446,0.0012150385,0.008997003],"genre_scores_gemma":[0.63621974,0.00063880865,0.35845485,0.0001228751,0.00007589453,0.00058318133,0.0028628602,0.00007669262,0.00096512993],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9868706,0.0043498375,0.002385074,0.0017136167,0.004249503,0.00043136597],"domain_scores_gemma":[0.94003415,0.049412824,0.002577809,0.00313623,0.004339576,0.00049940305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0103999525,0.00055240444,0.0017055975,0.016045822,0.0012435325,0.003917448,0.0017037943,0.0016873048,0.0011304655],"category_scores_gemma":[0.05728283,0.0004123004,0.0016526869,0.010818433,0.0011347012,0.005620047,0.0021265603,0.0014447839,0.0004925635],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001209559,0.0016014413,0.058965378,0.001156117,0.0006115202,0.0014246707,0.0036814832,0.024971882,0.009092537,0.02471643,0.0032438932,0.8693251],"study_design_scores_gemma":[0.00023887865,0.0009958732,0.05975556,0.0009619266,0.0011356176,0.0041900766,0.0039655133,0.7528798,0.025090879,0.13656238,0.013991401,0.0002320227],"about_ca_topic_score_codex":0.0026697118,"about_ca_topic_score_gemma":0.0018890558,"teacher_disagreement_score":0.016045822,"about_ca_system_score_codex":0.0009877175,"about_ca_system_score_gemma":0.0013191592,"threshold_uncertainty_score":0.05500084},"labels":[],"label_agreement":null},{"id":"W2004758929","doi":"10.1016/j.sysarc.2010.06.003","title":"Using complexity, coupling, and cohesion metrics as early indicators of vulnerabilities","year":2010,"lang":"en","type":"article","venue":"Journal of Systems Architecture","topic":"Software Engineering Research","field":"Computer Science","cited_by":289,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Cohesion (chemistry); Software security assurance; Software quality; Secure coding; Software metric; Software bug; Vulnerability (computing); Software; Data mining; Software evolution; Vulnerability management; Software development; Software engineering; Computer security; Software construction; Vulnerability assessment; Information security","score_opus":0.029593890437432966,"score_gpt":0.29178655686711547,"score_spread":0.2621926664296825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2004758929","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93767756,0.00039156017,0.056282196,0.0002562491,0.000059010796,0.00011295941,0.0006901937,0.0006088943,0.0039213514],"genre_scores_gemma":[0.9728212,0.00010296834,0.025647242,0.000015805304,0.000012672541,0.000057367433,0.00051059405,0.000046631212,0.0007855836],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99603003,0.0009431827,0.00035785456,0.0003128943,0.0020853495,0.00027062197],"domain_scores_gemma":[0.9644659,0.017670646,0.0074219885,0.0019247987,0.006774103,0.001742611],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038984944,0.0012279833,0.0005206907,0.009127379,0.00053162494,0.0021250339,0.0005913723,0.00093681604,0.0010505259],"category_scores_gemma":[0.036328115,0.00035091382,0.0007630214,0.0037824106,0.000569847,0.0039420635,0.001363204,0.0011977666,0.000358731],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065235497,0.00046223722,0.7759145,0.00024557146,0.0004998877,0.00027610344,0.0013766966,0.018179365,0.020157175,0.0050293063,0.0022430392,0.17496373],"study_design_scores_gemma":[0.00007200924,0.0021137835,0.6917934,0.00020557332,0.00061550346,0.0012050605,0.0015786368,0.24709748,0.03518658,0.0151247475,0.004734305,0.00027289186],"about_ca_topic_score_codex":0.002795093,"about_ca_topic_score_gemma":0.0068085473,"teacher_disagreement_score":0.009127379,"about_ca_system_score_codex":0.0007556871,"about_ca_system_score_gemma":0.0011666329,"threshold_uncertainty_score":0.020617485},"labels":[],"label_agreement":null},{"id":"W2005497056","doi":"10.1109/icsm.2010.5609670","title":"Software process recovery using Recovered Unified Process Views","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Process (computing); Software engineering; Software development process; Software; Goal-Driven Software Development Process; Software development; Software bug; BitTorrent tracker; Data mining; Artificial intelligence; Programming language","score_opus":0.03949954530939974,"score_gpt":0.315250383825109,"score_spread":0.27575083851570925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2005497056","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016395006,0.0001516777,0.9788456,0.00009639547,0.000017022765,0.00012368597,0.0003986176,0.003350957,0.00062105554],"genre_scores_gemma":[0.18136033,0.0002584025,0.81271976,0.000055133874,0.00003142632,0.00028618658,0.0033524975,0.0006921038,0.0012441679],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99329585,0.001745743,0.00053518225,0.0015498104,0.002468183,0.00040523364],"domain_scores_gemma":[0.9801334,0.005009792,0.0028184718,0.007672134,0.004082878,0.000283306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044419076,0.0011558265,0.0012809333,0.0064402195,0.00084844406,0.0033635164,0.0020927228,0.0015672037,0.0011483151],"category_scores_gemma":[0.028594624,0.00082458975,0.002293972,0.004609364,0.0010122075,0.0049646185,0.0028315978,0.002477516,0.0009846799],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055128476,0.00033085115,0.014343006,0.0006129346,0.00031636102,0.0008106299,0.005366799,0.07588289,0.030957285,0.037451133,0.006687978,0.8266888],"study_design_scores_gemma":[0.000053654832,0.00021012283,0.008724501,0.00020529392,0.00023023556,0.0006773256,0.001352587,0.8552345,0.041044243,0.059129585,0.032963417,0.00017445115],"about_ca_topic_score_codex":0.005721272,"about_ca_topic_score_gemma":0.00609456,"teacher_disagreement_score":0.0064402195,"about_ca_system_score_codex":0.0012140105,"about_ca_system_score_gemma":0.0028628295,"threshold_uncertainty_score":0.023491323},"labels":[],"label_agreement":null},{"id":"W2005599223","doi":"10.1142/s1469026814500138","title":"SOFTWARE DEVELOPMENT EFFORT ESTIMATION USING CLASSICAL AND FUZZY ANALOGY: A CROSS-VALIDATION COMPARATIVE STUDY","year":2014,"lang":"en","type":"article","venue":"International Journal of Computational Intelligence and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Analogy; Computer science; Software; Benchmarking; Software development; Estimation; Fuzzy logic; Software sizing; Data mining; Artificial intelligence; Software engineering; Machine learning; Software construction; Systems engineering; Programming language; Linguistics","score_opus":0.0587535365511344,"score_gpt":0.39425473157658364,"score_spread":0.33550119502544923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2005599223","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9890295,0.0003516708,0.00900588,0.000028897854,0.000022511298,0.00010199516,0.00013855344,0.000036722813,0.0012842696],"genre_scores_gemma":[0.9925632,0.00010528933,0.006362895,0.000019651894,0.0000129023665,0.00010312778,0.00047967373,0.000015426922,0.00033781774],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98343533,0.011053306,0.00104788,0.0014282973,0.002695985,0.00033931356],"domain_scores_gemma":[0.8705302,0.09740353,0.0037420236,0.008828553,0.018342732,0.0011529939],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029784942,0.0006496672,0.0007870522,0.004055944,0.0007424063,0.00097616616,0.0014646553,0.0015656755,0.000872804],"category_scores_gemma":[0.071365155,0.0002886279,0.0009830793,0.0025349644,0.001285595,0.0016270083,0.00146642,0.00079374464,0.0003662181],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0073816716,0.014399245,0.50151086,0.0018221517,0.003211065,0.0006868807,0.012438824,0.12525566,0.014711751,0.007779991,0.004913892,0.305888],"study_design_scores_gemma":[0.00061557937,0.00969316,0.60880446,0.00044654318,0.0010711421,0.00070463435,0.005984382,0.3438543,0.014510006,0.0056015076,0.008489508,0.0002247841],"about_ca_topic_score_codex":0.0024378747,"about_ca_topic_score_gemma":0.0025677762,"teacher_disagreement_score":0.029784942,"about_ca_system_score_codex":0.0010450936,"about_ca_system_score_gemma":0.0007085658,"threshold_uncertainty_score":0.15751976},"labels":[],"label_agreement":null},{"id":"W2005680430","doi":"10.1109/tse.2014.2383381","title":"Range Fixes: Interactive Error Resolution for Software Configuration","year":2014,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Key Research and Development Program of China; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Range (aeronautics); Constraint (computer-aided design); Software; Simple (philosophy); Theoretical computer science; String (physics); Programming language; Mathematics","score_opus":0.019126007936626662,"score_gpt":0.26072867292940377,"score_spread":0.2416026649927771,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2005680430","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010434081,0.00034010224,0.92918706,0.00013777998,0.00006704961,0.00019813732,0.0007314586,0.05631746,0.0025868681],"genre_scores_gemma":[0.1372999,0.00020733863,0.8520986,0.00012922488,0.00003216625,0.00040731553,0.0018582453,0.006250191,0.0017170439],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9951283,0.001520052,0.00040178082,0.0009679686,0.0016970509,0.00028480584],"domain_scores_gemma":[0.9854344,0.009190155,0.0010786222,0.0033741905,0.0007235769,0.00019897852],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004688331,0.0025993215,0.0009901886,0.0040952493,0.0011017874,0.0023413284,0.0036185484,0.002163877,0.012403033],"category_scores_gemma":[0.027381612,0.0012792124,0.0017567931,0.0019188564,0.0022020463,0.0038595963,0.005157712,0.0021666433,0.003198738],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061468227,0.0002270844,0.008296364,0.0010877366,0.0002093676,0.00074405223,0.0020684307,0.09919992,0.015250532,0.029608607,0.037431985,0.80526125],"study_design_scores_gemma":[0.00026398647,0.00031019835,0.002606599,0.0004974145,0.00013931554,0.0013592592,0.0005889663,0.7919218,0.049242947,0.07973276,0.07305857,0.00027820002],"about_ca_topic_score_codex":0.0022414087,"about_ca_topic_score_gemma":0.0035525984,"teacher_disagreement_score":0.012403033,"about_ca_system_score_codex":0.00079638464,"about_ca_system_score_gemma":0.0011731504,"threshold_uncertainty_score":0.041492283},"labels":[],"label_agreement":null},{"id":"W2005989626","doi":"10.4236/jsea.2012.510092","title":"Software Measurement Methods: An Analysis of Two Designs","year":2012,"lang":"en","type":"article","venue":"Journal of Software Engineering and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Consistency (knowledge bases); Software; Computer science; Identification (biology); Software measurement; Data collection; Data mining; Reliability engineering; Software engineering; Software development; Software quality; Engineering; Artificial intelligence; Statistics; Mathematics; Programming language","score_opus":0.05872708588833401,"score_gpt":0.34278169605405734,"score_spread":0.2840546101657233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2005989626","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04591919,0.0042238333,0.91298616,0.001325681,0.0008081286,0.0078096194,0.0005122738,0.00049011555,0.025925005],"genre_scores_gemma":[0.16618533,0.0020380649,0.81584275,0.00050665543,0.0001929357,0.012550532,0.0002526007,0.00022564623,0.0022054194],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.7813421,0.15261754,0.010104241,0.008934915,0.045449514,0.0015516413],"domain_scores_gemma":[0.51869285,0.41714382,0.014415469,0.018464332,0.029901914,0.0013816995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11521156,0.0020923305,0.002379312,0.011377285,0.002140072,0.0073975394,0.0023497092,0.0026546384,0.007443859],"category_scores_gemma":[0.29528108,0.0013135399,0.002782036,0.0073025706,0.0064636488,0.0071867616,0.0046867426,0.002779841,0.0011914644],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013856386,0.00068699877,0.03220331,0.0067972355,0.0011966033,0.00034553543,0.019382529,0.0048099076,0.00337648,0.46095616,0.0034515604,0.46540806],"study_design_scores_gemma":[0.0019718034,0.0127910385,0.0786755,0.011278133,0.0026515785,0.0022364662,0.024363456,0.07029368,0.012056949,0.60717934,0.1756513,0.00085082446],"about_ca_topic_score_codex":0.0012571316,"about_ca_topic_score_gemma":0.00085402286,"teacher_disagreement_score":0.11521156,"about_ca_system_score_codex":0.0063031367,"about_ca_system_score_gemma":0.0073045064,"threshold_uncertainty_score":0.60930425},"labels":[],"label_agreement":null},{"id":"W2006192515","doi":"10.1109/wcre.2013.6671287","title":"The influence of non-technical factors on code review","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":96,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Process (computing); Code (set theory); Code review; Key (lock); Replicate; Variety (cybernetics); Source code; Component (thermodynamics); Empirical research; Software engineering; Static program analysis; Software; Software development; Computer security; Artificial intelligence; Operating system; Programming language","score_opus":0.01925797591065593,"score_gpt":0.2970629541509352,"score_spread":0.2778049782402793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2006192515","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97796977,0.00078191655,0.0078024347,0.0011975337,0.00011472234,0.0007566105,0.0001658662,0.00032975958,0.0108814025],"genre_scores_gemma":[0.9933561,0.00018517987,0.0042167995,0.00021651413,0.00007437268,0.00028666726,0.0001051186,0.000098513694,0.0014608017],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.79901093,0.1331118,0.018709376,0.011158005,0.03294078,0.0050690183],"domain_scores_gemma":[0.06846052,0.80940783,0.07280342,0.015334816,0.027679577,0.0063138413],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09492455,0.0006652422,0.0008147416,0.004429192,0.0026988382,0.007137925,0.0017623244,0.0015076826,0.004219257],"category_scores_gemma":[0.58430696,0.0007399487,0.0009426653,0.0031054781,0.0025417723,0.0032976388,0.0023155536,0.0019080507,0.0014440495],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027739948,0.002473904,0.768885,0.0017775759,0.00062406465,0.00083429314,0.030144375,0.0028378074,0.008055654,0.001504909,0.0028369022,0.17725143],"study_design_scores_gemma":[0.00015292599,0.0021133614,0.9712063,0.00029673483,0.00022422527,0.0005528976,0.009417993,0.0051305247,0.0036731598,0.0014736415,0.0055764792,0.00018170786],"about_ca_topic_score_codex":0.0046205684,"about_ca_topic_score_gemma":0.006869078,"teacher_disagreement_score":0.90507543,"about_ca_system_score_codex":0.0039193816,"about_ca_system_score_gemma":0.007506562,"threshold_uncertainty_score":0.502015},"labels":[],"label_agreement":null},{"id":"W2006202761","doi":"10.1145/1984708.1984722","title":"Fishtail","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Plug-in; Programmer; Computer science; Software engineering; Documentation; Eclipse; Debugging; Task (project management); Java; Reuse; Interface (matter); Microsoft Visual Studio; Software; World Wide Web; Operating system; Engineering; Systems engineering","score_opus":0.0475835964734028,"score_gpt":0.24287993015590428,"score_spread":0.1952963336825015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2006202761","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.082315505,0.0016391091,0.12485999,0.0032032174,0.0019241421,0.0008146767,0.019618573,0.10613308,0.6594916],"genre_scores_gemma":[0.11713332,0.0010335668,0.09671416,0.002913641,0.0001523214,0.00063532894,0.03416248,0.015322939,0.7319322],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99941874,0.000030334419,0.000030542695,0.00018278632,0.00026787896,0.00006979601],"domain_scores_gemma":[0.99897516,0.00012446466,0.0000716204,0.00024856935,0.0003921502,0.00018806332],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00077451405,0.000735733,0.00058447133,0.0012996813,0.0011516288,0.0017547823,0.0019505528,0.0011303164,0.113288336],"category_scores_gemma":[0.0016821352,0.000438051,0.0007838084,0.0009669302,0.00057532056,0.0026625798,0.0025577073,0.0010697376,0.08520989],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010979678,0.00029051362,0.01367804,0.0008056819,0.00008125467,0.0013866135,0.0012406004,0.0008541545,0.11077243,0.020471442,0.365411,0.48391017],"study_design_scores_gemma":[0.000055233428,0.00016648554,0.006326203,0.000067828085,0.00003473974,0.000632411,0.000121650184,0.0012095366,0.0073436205,0.0023155508,0.9816834,0.000043379765],"about_ca_topic_score_codex":0.0069284244,"about_ca_topic_score_gemma":0.013364237,"teacher_disagreement_score":0.88671166,"about_ca_system_score_codex":0.00091837946,"about_ca_system_score_gemma":0.0011630673,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2006552151","doi":"10.1109/vissof.2007.4290705","title":"Visualization Patterns: A Context-Sensitive Tool to Evaluate Visualization Techniques","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada; Concordia University","keywords":"Visualization; Computer science; Context (archaeology); Software visualization; Data visualization; Software; Information visualization; Human–computer interaction; Visual analytics; Creative visualization; Representation (politics); Data science; Data mining; Software system; Programming language; Component-based software engineering","score_opus":0.025133364834826918,"score_gpt":0.3548940557982456,"score_spread":0.3297606909634187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2006552151","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15689497,0.00073088985,0.7568583,0.001607156,0.00040308235,0.004751503,0.008333797,0.056213785,0.014206539],"genre_scores_gemma":[0.25882807,0.00026954873,0.7290798,0.00025476422,0.00009998445,0.0037336245,0.0029937525,0.0026126832,0.00212782],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98495734,0.0074642766,0.0027336082,0.0008609966,0.0035931114,0.00039061348],"domain_scores_gemma":[0.8721762,0.09575555,0.007818135,0.011037436,0.01147238,0.001740301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017317055,0.0026809678,0.0016267326,0.01082869,0.0011737405,0.0051761633,0.0018485702,0.0028565775,0.006950434],"category_scores_gemma":[0.09845965,0.0010129528,0.0012552321,0.0052681165,0.0009871057,0.007898379,0.0035072728,0.0023158907,0.0015343787],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005957557,0.0026138714,0.05235099,0.0046856045,0.001147856,0.0012471605,0.014681296,0.01369527,0.056594785,0.033391584,0.05927293,0.75436103],"study_design_scores_gemma":[0.0031040057,0.009711986,0.09807835,0.0029628628,0.0017828934,0.0037579339,0.012811358,0.44936436,0.13917856,0.10655728,0.17094049,0.0017500032],"about_ca_topic_score_codex":0.0009852017,"about_ca_topic_score_gemma":0.001347042,"teacher_disagreement_score":0.017317055,"about_ca_system_score_codex":0.00087383215,"about_ca_system_score_gemma":0.001103191,"threshold_uncertainty_score":0.09158242},"labels":[],"label_agreement":null},{"id":"W2006700268","doi":"10.1145/1117696.1117704","title":"Coping with an open bug repository","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":248,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Software bug; Eclipse; Open source; Software engineering; World Wide Web; Software; Process (computing); Java; Software development; Security bug; Android (operating system); Data science; Operating system; Cloud computing","score_opus":0.02219749091500454,"score_gpt":0.28874406658755963,"score_spread":0.2665465756725551,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2006700268","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48144412,0.002310635,0.466573,0.02009029,0.0007898261,0.0006004192,0.00031728734,0.008864376,0.019009996],"genre_scores_gemma":[0.770757,0.0010661463,0.20968145,0.0027284275,0.00042875865,0.00049803185,0.0007876849,0.0019649193,0.012087618],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9633021,0.01280484,0.0032229053,0.0043923585,0.013873617,0.0024041587],"domain_scores_gemma":[0.80441535,0.08640212,0.04559107,0.03501748,0.020469889,0.008104151],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022738257,0.0013997382,0.0012999976,0.006381666,0.0050785695,0.008226728,0.0054118973,0.0047206916,0.00306972],"category_scores_gemma":[0.15053459,0.0018793565,0.0014121581,0.004773589,0.004784863,0.01987952,0.011124378,0.006086729,0.0012157607],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069718016,0.001223223,0.16266373,0.0015894577,0.0004091095,0.016479349,0.0731084,0.021333363,0.015426964,0.08156796,0.051655646,0.5738457],"study_design_scores_gemma":[0.00041197246,0.0020285635,0.07166822,0.0030199634,0.001003464,0.06937035,0.053724933,0.18528298,0.04192371,0.2383873,0.33163354,0.0015450345],"about_ca_topic_score_codex":0.0027365438,"about_ca_topic_score_gemma":0.002969519,"teacher_disagreement_score":0.022738257,"about_ca_system_score_codex":0.0026535504,"about_ca_system_score_gemma":0.0042584753,"threshold_uncertainty_score":0.12025285},"labels":[],"label_agreement":null},{"id":"W2006870009","doi":"10.1145/1808920.1808934","title":"What is trust in a recommender for software development?","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Recommender system; Computer science; Position paper; Software; Software development; World Wide Web; Position (finance); Software engineering; Knowledge management; Business","score_opus":0.0272711915748344,"score_gpt":0.2916732873601709,"score_spread":0.2644020957853365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2006870009","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6780911,0.008188853,0.15611705,0.08195592,0.00062269944,0.0002598341,0.00041385044,0.00045465087,0.073896125],"genre_scores_gemma":[0.9901669,0.00079632155,0.007283295,0.00035889776,0.000106411484,0.000019357323,0.00005001972,0.000020344345,0.0011984293],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.98971784,0.005405639,0.00084456307,0.00096003606,0.002262482,0.00080944004],"domain_scores_gemma":[0.8977989,0.06333645,0.013397003,0.0060772207,0.015049358,0.004341119],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012754995,0.00035365127,0.0011243541,0.0017561067,0.0024343878,0.006722057,0.0010306402,0.0037346429,0.0029052442],"category_scores_gemma":[0.11666581,0.00064706086,0.0006669276,0.001983066,0.0029711963,0.013456598,0.0015429998,0.0024312944,0.00075467135],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000936925,0.00088842626,0.38590035,0.0016854472,0.0012312335,0.0014628533,0.020028155,0.01531779,0.0042650173,0.1598332,0.017641138,0.3908094],"study_design_scores_gemma":[0.00033344227,0.0018468061,0.21374275,0.0015421456,0.0014076805,0.002524332,0.029232118,0.19796392,0.0060022753,0.48431897,0.06038824,0.00069739553],"about_ca_topic_score_codex":0.017258119,"about_ca_topic_score_gemma":0.011395927,"teacher_disagreement_score":0.017258119,"about_ca_system_score_codex":0.0039965827,"about_ca_system_score_gemma":0.0022056543,"threshold_uncertainty_score":0.06745565},"labels":[],"label_agreement":null},{"id":"W2007210734","doi":"10.1109/msr.2013.6624016","title":"Will my patch make it? And how fast? Case study on the Linux kernel","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":130,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; Polytechnique Montréal","funders":"","keywords":"Linux kernel; Computer science; Kernel (algebra); Process (computing); Operating system; World Wide Web; Control (management); Software engineering; Artificial intelligence; Mathematics","score_opus":0.028534250511539593,"score_gpt":0.26675876547427196,"score_spread":0.23822451496273236,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007210734","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99175245,0.00092732545,0.0016493805,0.0013057141,0.000052577467,0.000087819135,0.00011925886,0.000084700856,0.004020751],"genre_scores_gemma":[0.99144274,0.0008080058,0.0035397494,0.0003658808,0.00008708443,0.0000713483,0.00016103692,0.00010754497,0.0034166202],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9910904,0.004774698,0.0006671844,0.00072353147,0.0021535703,0.000590576],"domain_scores_gemma":[0.8823704,0.082076594,0.01199583,0.0036128373,0.015166189,0.0047781775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009773417,0.00041536367,0.00044939472,0.0028219405,0.0033349618,0.002149031,0.0012287316,0.0017942274,0.0014967374],"category_scores_gemma":[0.056959692,0.00039850228,0.00039282723,0.0028847668,0.0012481516,0.0029310947,0.0011491495,0.0015056668,0.0006673911],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010125646,0.0029000903,0.35438034,0.0018260349,0.00019480305,0.059174143,0.31228906,0.00412298,0.012224279,0.0038227856,0.031168453,0.21688446],"study_design_scores_gemma":[0.0001640649,0.0029922712,0.47567868,0.0011103807,0.00036778825,0.03516582,0.29047742,0.012417099,0.013495043,0.0024297189,0.16531114,0.0003905144],"about_ca_topic_score_codex":0.008401315,"about_ca_topic_score_gemma":0.017216738,"teacher_disagreement_score":0.009773417,"about_ca_system_score_codex":0.0019990115,"about_ca_system_score_gemma":0.0013248791,"threshold_uncertainty_score":0.05168742},"labels":[],"label_agreement":null},{"id":"W2007234571","doi":"10.1631/jzus.c1300102","title":"An experimental study on the conversion between IFPUG and UCP functional size measurement units","year":2014,"lang":"en","type":"article","venue":"Journal of Zhejiang University SCIENCE C","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Software; Function point; Measure (data warehouse); Function (biology); Point (geometry); Functional requirement; Software measurement; Data mining; Software development; Mathematics; Software engineering; Component-based software engineering; Programming language","score_opus":0.06264083403630134,"score_gpt":0.2566534389328208,"score_spread":0.19401260489651945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007234571","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96017134,0.00030211097,0.027627518,0.000104291015,0.00028207866,0.0011331575,0.00044021825,0.00030583455,0.009633494],"genre_scores_gemma":[0.94984925,0.00029158444,0.040999077,0.000105634994,0.00006098416,0.0018553857,0.0005388818,0.00014303022,0.006156079],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9919023,0.0031979822,0.0007162,0.0017120112,0.002120679,0.00035082025],"domain_scores_gemma":[0.93842685,0.04406971,0.003027231,0.0057518864,0.007844857,0.00087956194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007513948,0.00079502433,0.00055065256,0.0013906275,0.0005934428,0.0010128287,0.0011748609,0.0008316792,0.012021534],"category_scores_gemma":[0.048672616,0.0004112586,0.00050185836,0.0012259744,0.0011020314,0.0015172095,0.0011064776,0.00110021,0.0017316582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.021328509,0.025213854,0.08883256,0.0024304562,0.00022356295,0.00067653664,0.013811353,0.003833371,0.24497348,0.012799796,0.0048302975,0.5810463],"study_design_scores_gemma":[0.0012989378,0.13953878,0.42316604,0.0008291669,0.00075595744,0.0015002121,0.01463209,0.049488783,0.32609865,0.009481872,0.032756686,0.00045294187],"about_ca_topic_score_codex":0.0005727185,"about_ca_topic_score_gemma":0.0005103671,"teacher_disagreement_score":0.012021534,"about_ca_system_score_codex":0.00044740946,"about_ca_system_score_gemma":0.00058558513,"threshold_uncertainty_score":0.04021609},"labels":[],"label_agreement":null},{"id":"W2007293407","doi":"10.5555/2666527.2666530","title":"Clone detection meets semantic web-based transitive closure computation","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Source code; Transitive relation; Transitive closure; Semantic Web; Programming language; Theoretical computer science; Data mining; Artificial intelligence; Mathematics","score_opus":0.014909508081670238,"score_gpt":0.25853503482566575,"score_spread":0.2436255267439955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007293407","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048535958,0.00011599831,0.94439167,0.00042987167,0.000042610925,0.00022400449,0.00020504901,0.0020754451,0.003979369],"genre_scores_gemma":[0.54822147,0.00016275763,0.44811028,0.00019718331,0.00009648673,0.00031031112,0.00084517593,0.0003686463,0.0016877044],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98963296,0.0022024887,0.00084072823,0.002619159,0.004055239,0.0006495003],"domain_scores_gemma":[0.96702284,0.020884747,0.002789105,0.0048354063,0.003952299,0.00051550666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005155579,0.0008336508,0.0015274704,0.0046693645,0.0018875342,0.004421458,0.0018256233,0.0014561836,0.0022065917],"category_scores_gemma":[0.03786172,0.00066051324,0.0031133266,0.0020842191,0.004013483,0.00847498,0.0035148903,0.0019325734,0.00057549635],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064038835,0.0005213376,0.012227824,0.0007386027,0.00029205784,0.0016860222,0.0025624367,0.08576128,0.027978487,0.53574014,0.004599097,0.3272523],"study_design_scores_gemma":[0.000066104236,0.00012035279,0.0012860216,0.00008274432,0.00013350419,0.0005291421,0.0004025154,0.502102,0.026946504,0.4627632,0.005501851,0.00006600332],"about_ca_topic_score_codex":0.005114637,"about_ca_topic_score_gemma":0.0040157554,"teacher_disagreement_score":0.005155579,"about_ca_system_score_codex":0.002286614,"about_ca_system_score_gemma":0.0028059965,"threshold_uncertainty_score":0.027265608},"labels":[],"label_agreement":null},{"id":"W2007351432","doi":"10.5555/2664398.2664399","title":"An accurate estimation of the Levenshtein distance using metric trees and Manhattan distance","year":2012,"lang":"en","type":"article","venue":"International Workshop on Software Clones","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Levenshtein distance; Metric (unit); Precision and recall; Computer science; Edit distance; Software; Euclidean distance; Distance measurement; Data mining; Artificial intelligence; Engineering","score_opus":0.036183713733103924,"score_gpt":0.32937564789469387,"score_spread":0.29319193416158995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007351432","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01880349,0.00084059144,0.9784676,0.000052741383,0.000050650997,0.000029995646,0.00007707534,0.00084317836,0.00083458825],"genre_scores_gemma":[0.24828522,0.0008425491,0.7477258,0.00004911811,0.00009746249,0.000088058834,0.0005384202,0.00027625554,0.0020971335],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99553204,0.0008334024,0.00026109375,0.00070819235,0.0024568704,0.00020839629],"domain_scores_gemma":[0.99178773,0.003510188,0.0010029149,0.0010554992,0.0024649913,0.00017864046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001897747,0.0010370066,0.0011136916,0.006246792,0.00075731205,0.0016939541,0.001440268,0.001056086,0.0010825295],"category_scores_gemma":[0.015847249,0.00052391127,0.0007026815,0.0040867208,0.0008282042,0.00418226,0.0012333379,0.0011236881,0.0012363879],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027693153,0.00012512639,0.01718274,0.00043351017,0.0002726095,0.0003689789,0.00076922344,0.07282839,0.052381825,0.037316922,0.004241253,0.8138024],"study_design_scores_gemma":[0.00002332282,0.00039397556,0.01518772,0.0000916693,0.00009338394,0.0020061524,0.00035068326,0.8623321,0.05597183,0.03966514,0.023661619,0.00022236939],"about_ca_topic_score_codex":0.0034338178,"about_ca_topic_score_gemma":0.0032215393,"teacher_disagreement_score":0.006246792,"about_ca_system_score_codex":0.00095459475,"about_ca_system_score_gemma":0.000790544,"threshold_uncertainty_score":0.010036409},"labels":[],"label_agreement":null},{"id":"W2007425631","doi":"10.1109/scam.2013.6648192","title":"JSNOSE: Detecting JavaScript Code Smells","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":139,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"JavaScript; Computer science; Code refactoring; Code smell; Scripting language; Program comprehension; Programming language; Static program analysis; Web application; Program slicing; Source code; Unobtrusive JavaScript; Code (set theory); World Wide Web; Software engineering; Software quality; Set (abstract data type); Software; Software development; Software system; Rich Internet application","score_opus":0.019232873970876637,"score_gpt":0.25036886794030394,"score_spread":0.2311359939694273,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007425631","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86834145,0.000654455,0.09889697,0.00018065453,0.000072805335,0.00055095286,0.0058296733,0.02241379,0.0030593434],"genre_scores_gemma":[0.82575965,0.00031115522,0.16113862,0.00009735913,0.000033185104,0.00033253836,0.00838893,0.0010498434,0.0028887703],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962405,0.00039777547,0.0004502081,0.0005443896,0.0022058883,0.00016123962],"domain_scores_gemma":[0.982971,0.0048896107,0.005157548,0.0015622056,0.0046956916,0.00072391506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016151728,0.0010071751,0.00068740716,0.006079821,0.00037409132,0.0009271419,0.00075943244,0.0007070987,0.00063417596],"category_scores_gemma":[0.013174218,0.00034476625,0.00047297883,0.0023321905,0.0003554364,0.0012497492,0.0011370904,0.0005558644,0.0006227945],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000734055,0.00085657573,0.39704236,0.0018336804,0.00041912758,0.001412003,0.0026357255,0.007132637,0.16318248,0.00085866713,0.012674232,0.4112184],"study_design_scores_gemma":[0.000096238844,0.0011257885,0.5768176,0.00018796632,0.0002153704,0.002785607,0.0011437199,0.22956182,0.1691125,0.0014121534,0.017269459,0.000271861],"about_ca_topic_score_codex":0.0023397596,"about_ca_topic_score_gemma":0.00547837,"teacher_disagreement_score":0.006079821,"about_ca_system_score_codex":0.0003934085,"about_ca_system_score_gemma":0.0006861263,"threshold_uncertainty_score":0.008541942},"labels":[],"label_agreement":null},{"id":"W2007472189","doi":"10.1007/s10009-009-0122-5","title":"An approach for estimating the time needed to perform code changes in business applications","year":2009,"lang":"en","type":"article","venue":"International Journal on Software Tools for Technology Transfer","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Business process modeling; Artifact-centric business process model; Computer science; Business Process Model and Notation; Business rule; Business process management; Business process discovery; Business process; Source code; Metric (unit); Competitive advantage; Process management; Task (project management); Business; Work in process; Programming language; Systems engineering; Marketing","score_opus":0.02550414897938195,"score_gpt":0.30783355710772803,"score_spread":0.28232940812834606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007472189","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21133998,0.0012591172,0.7803951,0.0002028318,0.00009847463,0.00024356651,0.00075806736,0.0024837174,0.003219045],"genre_scores_gemma":[0.5468384,0.0003588564,0.44946307,0.000044938042,0.00004455016,0.00016960556,0.00062218367,0.00017446287,0.0022839026],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99782264,0.00035594593,0.00015689508,0.0005057251,0.0009757808,0.00018304902],"domain_scores_gemma":[0.98996896,0.006378615,0.0011670474,0.00065532734,0.0015526568,0.00027739772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016654981,0.0013130521,0.0009255564,0.004580856,0.00074433745,0.0015721994,0.00137645,0.0016569289,0.0020827071],"category_scores_gemma":[0.01375511,0.00073224446,0.0010490733,0.0029802236,0.00044361473,0.001849416,0.00085845834,0.0010464655,0.0006650685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010850225,0.0005262064,0.046739623,0.00045403093,0.0003082922,0.00026153715,0.00055895676,0.30988964,0.039113183,0.006959275,0.0019952133,0.5921091],"study_design_scores_gemma":[0.00003616323,0.00029415154,0.011250418,0.000024078607,0.00010131914,0.00016848344,0.00015236504,0.9717665,0.010575262,0.0037791464,0.0018015187,0.000050583803],"about_ca_topic_score_codex":0.021168666,"about_ca_topic_score_gemma":0.018609144,"teacher_disagreement_score":0.021168666,"about_ca_system_score_codex":0.0016907662,"about_ca_system_score_gemma":0.0018446811,"threshold_uncertainty_score":0.042090952},"labels":[],"label_agreement":null},{"id":"W2007818470","doi":"10.5555/2662708.2662710","title":"A mutation analysis based benchmarking framework for clone detectors","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Saskatchewan","funders":"","keywords":"Benchmarking; Benchmark (surveying); Granularity; clone (Java method); Computer science; Detector; Mutation testing; Mutation; Data mining; Software engineering; Operating system; Biology; Genetics; Telecommunications; Business; Gene; Cartography; Geography","score_opus":0.015045571800584745,"score_gpt":0.2753145520658611,"score_spread":0.26026898026527634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007818470","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0403532,0.0005063274,0.93802273,0.00025595768,0.000058279373,0.00037960452,0.00031252185,0.017431892,0.0026794225],"genre_scores_gemma":[0.45251766,0.00018669337,0.543695,0.00014124985,0.00004141095,0.0004200601,0.00097402977,0.001154176,0.00086981634],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9791915,0.009357635,0.0016779122,0.002132505,0.0066907285,0.0009497524],"domain_scores_gemma":[0.9712669,0.011942809,0.0034301162,0.00498488,0.0074616973,0.00091360445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019387884,0.0020049335,0.0020350427,0.008636074,0.0010566899,0.0037256232,0.004716104,0.002430047,0.0019472713],"category_scores_gemma":[0.043343995,0.00070602197,0.0013570382,0.0036037145,0.0016185482,0.0037856768,0.0026642075,0.0019050954,0.0006259677],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008950399,0.0013702771,0.027694618,0.00082754414,0.00044296178,0.0007916025,0.0007496164,0.3002535,0.050595473,0.08320578,0.009201401,0.52397215],"study_design_scores_gemma":[0.00003886331,0.00029980487,0.002218698,0.00008833821,0.000052746087,0.00022533374,0.00006272684,0.9619203,0.016053537,0.015347665,0.00362596,0.000065952416],"about_ca_topic_score_codex":0.004795802,"about_ca_topic_score_gemma":0.0032293394,"teacher_disagreement_score":0.019387884,"about_ca_system_score_codex":0.0025014635,"about_ca_system_score_gemma":0.0032014179,"threshold_uncertainty_score":0.102534175},"labels":[],"label_agreement":null},{"id":"W2008063114","doi":"10.1145/1882291.1882335","title":"DSketch","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Dependency (UML); Software engineering; Software; Process (computing); Software maintenance; Software system; Programming language; Software development","score_opus":0.01029217490388157,"score_gpt":0.25827335784387734,"score_spread":0.24798118293999577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2008063114","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009374238,0.0010112411,0.36745992,0.0009073983,0.0010006863,0.00062611134,0.026547741,0.5073571,0.08571555],"genre_scores_gemma":[0.1205065,0.0020782705,0.42539856,0.001621284,0.00027404775,0.0019232716,0.10691014,0.11893072,0.22235717],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99784327,0.00032800587,0.00022757839,0.0004965013,0.00091471133,0.0001898691],"domain_scores_gemma":[0.9949366,0.0020547053,0.00022798333,0.0014428337,0.0010800288,0.00025788258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021449071,0.0013022951,0.0009965317,0.0029759181,0.00091222994,0.0025852176,0.0026453983,0.0013285151,0.09735528],"category_scores_gemma":[0.0105416225,0.001418439,0.0013568714,0.0014795638,0.00082356663,0.0050905147,0.003984699,0.0029451293,0.04829704],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009215254,0.00021023945,0.0034968872,0.0016964946,0.00014351305,0.00074691017,0.0008515271,0.006428528,0.009619017,0.04839459,0.49844027,0.42905053],"study_design_scores_gemma":[0.00027400715,0.00012490992,0.0015460064,0.00025424286,0.00006309122,0.0010122219,0.00015988246,0.025106288,0.017804515,0.026341753,0.927168,0.00014499354],"about_ca_topic_score_codex":0.0027458614,"about_ca_topic_score_gemma":0.004017463,"teacher_disagreement_score":0.09735528,"about_ca_system_score_codex":0.0009887208,"about_ca_system_score_gemma":0.0021969143,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2008107570","doi":"10.1109/ms.2009.193","title":"What Makes APIs Hard to Learn? Answers from Developers","year":2009,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":383,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Microsoft Research","keywords":"Application programming interface; Computer science; Learnability; Usability; Reuse; Software engineering; World Wide Web; Interface (matter); Software development; Software; Human–computer interaction; Programming language; Engineering; Operating system","score_opus":0.023596069061703036,"score_gpt":0.26713194313129746,"score_spread":0.24353587406959443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2008107570","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013295476,0.010994925,0.003063412,0.94804657,0.0022495324,0.000024352395,0.000054405235,0.00015643646,0.022114914],"genre_scores_gemma":[0.51415807,0.045967102,0.011075129,0.39016768,0.005548192,0.000289885,0.00030125008,0.000632512,0.03186024],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.98232687,0.008379487,0.0008673149,0.0013716335,0.004410093,0.0026446315],"domain_scores_gemma":[0.89213806,0.060652576,0.0072354055,0.0039958167,0.0224439,0.013534249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019402629,0.00095182116,0.0009812808,0.0024137006,0.0085157985,0.011825908,0.0019542715,0.013561422,0.009096863],"category_scores_gemma":[0.10662015,0.0009219491,0.00081655796,0.002150235,0.011404638,0.026911188,0.009880471,0.013828562,0.0036093707],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007009907,0.00023396483,0.017367352,0.0007834042,0.000056886107,0.0016459968,0.09351159,0.00017438867,0.00047662618,0.05148369,0.60070765,0.23348835],"study_design_scores_gemma":[0.00006524231,0.00009706218,0.00606483,0.002136572,0.000057778234,0.0023352199,0.29969138,0.00044996294,0.00041165965,0.08879763,0.59975106,0.00014159766],"about_ca_topic_score_codex":0.0061043715,"about_ca_topic_score_gemma":0.008344947,"teacher_disagreement_score":0.019402629,"about_ca_system_score_codex":0.004103102,"about_ca_system_score_gemma":0.008270158,"threshold_uncertainty_score":0.10261208},"labels":[],"label_agreement":null},{"id":"W2008174757","doi":"10.1109/qsic.2014.46","title":"A Comparative Study of Invariants Generated by Daikon and User-Defined Design Contracts","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Haystack; Computer science; Quality (philosophy); Set (abstract data type); Complement (music); Control (management); Programming language; Artificial intelligence","score_opus":0.052250477998537034,"score_gpt":0.28860942132893574,"score_spread":0.2363589433303987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2008174757","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8929986,0.00043178277,0.09419027,0.00015643566,0.00003245714,0.00015137803,0.0006822749,0.008698307,0.002658616],"genre_scores_gemma":[0.89260805,0.00017154022,0.10288332,0.000070240756,0.00000714389,0.00011256342,0.0020954323,0.0011479974,0.0009037064],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99036777,0.0027825562,0.0009333436,0.0011502438,0.0041703847,0.00059567305],"domain_scores_gemma":[0.90913296,0.06424227,0.0073403926,0.013375812,0.005396105,0.000512507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008339338,0.0005775911,0.00038933987,0.0023026655,0.0003713849,0.0012693706,0.0012417194,0.0010418216,0.001490214],"category_scores_gemma":[0.053358406,0.0005350676,0.0007705384,0.0012590812,0.0011193826,0.0019814048,0.0012093442,0.0009626311,0.00024458888],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037329134,0.0015245539,0.16928147,0.0025607052,0.00045130344,0.001872415,0.006497558,0.17862919,0.1290723,0.019487195,0.0034657062,0.48342478],"study_design_scores_gemma":[0.00024667146,0.0026751764,0.0762667,0.0003220674,0.0002788419,0.001612524,0.0021349152,0.67739314,0.21352799,0.008054988,0.017278831,0.00020808539],"about_ca_topic_score_codex":0.0021273931,"about_ca_topic_score_gemma":0.0029073157,"teacher_disagreement_score":0.008339338,"about_ca_system_score_codex":0.0010178074,"about_ca_system_score_gemma":0.0012475335,"threshold_uncertainty_score":0.044103146},"labels":[],"label_agreement":null},{"id":"W2008190228","doi":"10.1109/iwsm-mensura.2011.52","title":"Design of a Functional Size Measurement Procedure for Real-Time Embedded Software Requirements Expressed using the Simulink Model","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Functional requirement; Software; Embedded software; Identification (biology); Software measurement; Embedded system; Software design; Reliability engineering; Software construction; Software system; Software engineering; Software development; Operating system; Engineering","score_opus":0.1965763316034351,"score_gpt":0.30020361007344765,"score_spread":0.10362727847001255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2008190228","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054188333,0.000008539259,0.9916648,0.000025423209,0.000006756605,0.00015810934,0.00004085003,0.0020038711,0.00067266735],"genre_scores_gemma":[0.14873382,0.00004663633,0.8486264,0.00004317724,0.000011922453,0.0006943656,0.00023801377,0.00038928044,0.0012163501],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9959637,0.001201528,0.00039439762,0.0006961943,0.0015989015,0.00014520372],"domain_scores_gemma":[0.992516,0.0028950374,0.000944887,0.0013208666,0.0022326612,0.000090630754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044005807,0.0007943193,0.0003725061,0.001708536,0.00036607706,0.0010194747,0.0010274848,0.00042398032,0.0021955643],"category_scores_gemma":[0.011792726,0.0004799572,0.0006068404,0.00041159103,0.00071855175,0.0011089953,0.0005073103,0.00090310775,0.0007363408],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000543665,0.0003901372,0.0061883368,0.00066049123,0.00014526249,0.00066554843,0.0018047651,0.21878304,0.18274713,0.09318912,0.0037462104,0.49113637],"study_design_scores_gemma":[0.000107936336,0.00051158573,0.0033861657,0.00012853034,0.00009148783,0.00042202824,0.00022001337,0.69867164,0.26054904,0.011170436,0.024610788,0.00013044976],"about_ca_topic_score_codex":0.0020101364,"about_ca_topic_score_gemma":0.001509504,"teacher_disagreement_score":0.0044005807,"about_ca_system_score_codex":0.0010689909,"about_ca_system_score_gemma":0.0026781426,"threshold_uncertainty_score":0.023272753},"labels":[],"label_agreement":null},{"id":"W2008377554","doi":"10.1109/qsic.2014.30","title":"Early Identification of Future Committers in Open Source Software Projects","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Japan Society for the Promotion of Science","keywords":"Commit; Promotion (chess); Identification (biology); Computer science; Eclipse; Permission; Software engineering; Software; Quality (philosophy); Code review; Source code; Software quality; Open source; Software development; Database; Operating system; Political science","score_opus":0.01696945354120046,"score_gpt":0.2672900264080182,"score_spread":0.25032057286681775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2008377554","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97792715,0.0011002455,0.016901575,0.0003713545,0.00003139588,0.00013229661,0.00014304537,0.00013048324,0.0032625976],"genre_scores_gemma":[0.9914684,0.0003184141,0.006557153,0.000041872565,0.000030769253,0.000051735755,0.00028692585,0.000022753436,0.001221933],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99394715,0.0017444055,0.0006050789,0.0011460806,0.0020634201,0.0004937295],"domain_scores_gemma":[0.84203273,0.08493189,0.04513698,0.0064727394,0.017088974,0.004336641],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008138697,0.00046547744,0.00051124254,0.005231695,0.0012627238,0.0017457578,0.0007627564,0.0010315006,0.0010739035],"category_scores_gemma":[0.07619309,0.0005186561,0.00036906995,0.001970417,0.0007606844,0.0031920625,0.0017780801,0.0011033104,0.00037277894],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020019992,0.00015336706,0.91009307,0.00020746821,0.000049566566,0.00073790294,0.0064237383,0.0012575318,0.0032944176,0.0018843317,0.00082202285,0.074876375],"study_design_scores_gemma":[0.000019326711,0.00023959637,0.9611797,0.00026388446,0.000076586904,0.001118953,0.005446434,0.018220283,0.0036898295,0.0042033675,0.0054707187,0.00007141279],"about_ca_topic_score_codex":0.004277108,"about_ca_topic_score_gemma":0.0070650433,"teacher_disagreement_score":0.008138697,"about_ca_system_score_codex":0.00088803034,"about_ca_system_score_gemma":0.0010861218,"threshold_uncertainty_score":0.043042064},"labels":[],"label_agreement":null},{"id":"W2009132449","doi":"10.1145/1268784.1268837","title":"Introducing students to professional software construction","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Key (lock); Software engineering; Software construction; Software; Software development; Personal software process; Social software engineering; Software peer review; Software analytics; Extreme programming practices; Engineering management; Engineering; Programming language","score_opus":0.011457741507453129,"score_gpt":0.31570610170425895,"score_spread":0.3042483601968058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2009132449","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5227892,0.005107021,0.18417756,0.032411523,0.0037517906,0.0023105056,0.0004389083,0.0036750005,0.24533848],"genre_scores_gemma":[0.62741786,0.0066991835,0.1631204,0.013342911,0.0012454635,0.0012226232,0.0005889488,0.0003591552,0.18600354],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986866,0.00023286638,0.000060298007,0.00018555814,0.000351714,0.00048299352],"domain_scores_gemma":[0.9956174,0.00063151674,0.00027117992,0.00023507324,0.00078299706,0.0024618614],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015479613,0.0008759116,0.0006112853,0.0006856256,0.0014975778,0.0036786522,0.0010556695,0.0018990854,0.01738951],"category_scores_gemma":[0.0043578963,0.00045060393,0.0006234674,0.00051489036,0.0012295177,0.001741728,0.0040170774,0.0026923595,0.005591712],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003544069,0.012960631,0.027757859,0.0012937569,0.00003075356,0.0021320663,0.04715778,0.0022547108,0.036194514,0.08302667,0.128196,0.65864074],"study_design_scores_gemma":[0.00014505787,0.0037848116,0.01271545,0.000806738,0.000039390547,0.0038481446,0.010350391,0.0030001749,0.017367624,0.03561478,0.91223156,0.00009582722],"about_ca_topic_score_codex":0.00089783117,"about_ca_topic_score_gemma":0.002126786,"teacher_disagreement_score":0.01738951,"about_ca_system_score_codex":0.00170451,"about_ca_system_score_gemma":0.003767588,"threshold_uncertainty_score":0.058173716},"labels":[],"label_agreement":null},{"id":"W2009151039","doi":"10.1016/j.jss.2012.07.050","title":"Towards an early software estimation using log-linear regression and a multilayer perceptron model","year":2012,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":214,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Use Case Points; Software; Computer science; Linear regression; Machine learning; Artificial neural network; Software sizing; Software development; Perceptron; Multilayer perceptron; Data mining; Software metric; Estimator; Artificial intelligence; Software development process; Software construction; Statistics; Mathematics","score_opus":0.050758474118636175,"score_gpt":0.3211822357308679,"score_spread":0.27042376161223175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2009151039","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009517129,0.00025827723,0.9887679,0.00016068427,0.000026489466,0.000011013163,0.000025169766,0.0007126056,0.0005206037],"genre_scores_gemma":[0.41515845,0.0006084739,0.57355434,0.00023070718,0.0001214359,0.000084942956,0.00025741445,0.00037005797,0.009614203],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986545,0.0004784479,0.00008810908,0.000271795,0.00036609397,0.00014104493],"domain_scores_gemma":[0.995368,0.0026737167,0.00030709087,0.00045784863,0.0010447543,0.00014863929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025531058,0.0009744114,0.0011627072,0.001414234,0.00047174827,0.0015989337,0.0018604527,0.0017208117,0.0019669554],"category_scores_gemma":[0.009591992,0.0011189667,0.0011465957,0.0011500452,0.00058225007,0.003522322,0.0016955425,0.003260963,0.001767915],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033637352,0.00035510404,0.005212521,0.00024952856,0.00019630487,0.00020995965,0.00023411056,0.3786913,0.019726537,0.022355508,0.0035868823,0.5688459],"study_design_scores_gemma":[0.0000030140004,0.000018815277,0.00030215958,0.0000095336,0.000011292918,0.000012563978,0.0000062128324,0.9931058,0.0016097237,0.0045319204,0.00038181525,0.0000071841873],"about_ca_topic_score_codex":0.007829616,"about_ca_topic_score_gemma":0.009257376,"teacher_disagreement_score":0.007829616,"about_ca_system_score_codex":0.0008466004,"about_ca_system_score_gemma":0.0013177112,"threshold_uncertainty_score":0.015568078},"labels":[],"label_agreement":null},{"id":"W2009204692","doi":"10.1016/j.infsof.2015.02.007","title":"Assessing the use of slicing-based visualizing techniques on the understanding of large metamodels","year":2015,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Metamodeling; Computer science; Slicing; Program slicing; Visualization; Correctness; Program comprehension; Human–computer interaction; Software engineering; Domain (mathematical analysis); Modeling language; Usability; Programming language; Artificial intelligence; Software; Software system; World Wide Web","score_opus":0.14777201499863313,"score_gpt":0.3510519464246652,"score_spread":0.20327993142603207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2009204692","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9359182,0.0003787357,0.058488872,0.0002126112,0.00001816001,0.00018216323,0.00032287522,0.0017469533,0.002731297],"genre_scores_gemma":[0.9082794,0.00025023086,0.09027129,0.000031517375,0.000007719247,0.000074990414,0.0004946065,0.00024881572,0.00034135897],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99149656,0.0047167838,0.0007396531,0.0006904221,0.0020408107,0.00031582313],"domain_scores_gemma":[0.7502254,0.20696504,0.011735767,0.014540522,0.014716296,0.0018169285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014588856,0.001350558,0.000553237,0.0045241327,0.00057246257,0.0029972657,0.0012447604,0.0017490666,0.001162344],"category_scores_gemma":[0.14581378,0.00058455043,0.0009638742,0.0022498993,0.00063257356,0.0055381786,0.0020271933,0.0014108724,0.00025466067],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003823297,0.0020271933,0.1607757,0.0020087434,0.0013208658,0.0005741039,0.025350355,0.09329972,0.07871675,0.00608891,0.0023696756,0.62364477],"study_design_scores_gemma":[0.00034505647,0.004635177,0.1546956,0.0007835236,0.001769263,0.00084169395,0.0092615085,0.7248854,0.08525414,0.009867526,0.0072676255,0.00039352794],"about_ca_topic_score_codex":0.004948425,"about_ca_topic_score_gemma":0.0066007217,"teacher_disagreement_score":0.014588856,"about_ca_system_score_codex":0.0008730282,"about_ca_system_score_gemma":0.0016432912,"threshold_uncertainty_score":0.07715422},"labels":[],"label_agreement":null},{"id":"W2009619860","doi":"10.1109/wcre.2013.6671314","title":"On the effect of program exploration on maintenance tasks","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Task (project management); Eclipse; Set (abstract data type); Duration (music); Software maintenance; Empirical research; Index (typography); Software engineering; Human–computer interaction; Software; Software development; World Wide Web; Programming language; Engineering; Systems engineering","score_opus":0.018485213930675242,"score_gpt":0.2730535047796239,"score_spread":0.25456829084894866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2009619860","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99431026,0.00041935418,0.003404296,0.00012389573,0.000009220255,0.000042564272,0.0002485295,0.0001999289,0.0012419249],"genre_scores_gemma":[0.99373674,0.00015487884,0.0047249976,0.000044304354,0.0000143012985,0.000055074805,0.0007326431,0.00008944708,0.00044754622],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99044496,0.0041530947,0.00088621787,0.0014864205,0.0024997983,0.0005295842],"domain_scores_gemma":[0.43997282,0.5133255,0.027688516,0.008281398,0.0072353357,0.0034963898],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010784114,0.00074578763,0.0005763686,0.0023619933,0.00072508527,0.0016582814,0.00071685103,0.00083217776,0.0016018753],"category_scores_gemma":[0.18604375,0.0005137584,0.000687324,0.0016955868,0.0010346122,0.0024405024,0.0012837449,0.0015388187,0.00033891352],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029931504,0.00090620335,0.82791966,0.0006605822,0.0005041165,0.00033681982,0.003314639,0.013631803,0.01455694,0.0003720559,0.0011536817,0.1336504],"study_design_scores_gemma":[0.00004790074,0.001201724,0.9739831,0.00005742226,0.00013976976,0.0002798392,0.0006636275,0.018662995,0.00362941,0.00040051184,0.0008760928,0.000057473593],"about_ca_topic_score_codex":0.0047877906,"about_ca_topic_score_gemma":0.007115868,"teacher_disagreement_score":0.010784114,"about_ca_system_score_codex":0.0010248095,"about_ca_system_score_gemma":0.000896481,"threshold_uncertainty_score":0.057032585},"labels":[],"label_agreement":null},{"id":"W2009798988","doi":"10.1109/compsac.2012.37","title":"Fine-Grained Design Pattern Detection","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Program comprehension; Eiffel; Software design pattern; Structural pattern; False positive paradox; Context (archaeology); Reverse engineering; Architectural pattern; Software design; Design pattern; Software; Field (mathematics); Programming language; Software system; Artificial intelligence; Software development; Object-oriented programming","score_opus":0.03425266214657047,"score_gpt":0.26047618227277497,"score_spread":0.2262235201262045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2009798988","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06077291,0.00031948145,0.92838776,0.00018484157,0.000058632802,0.00030086187,0.000508688,0.0073353513,0.0021315415],"genre_scores_gemma":[0.27919707,0.00020306327,0.7164095,0.00018377388,0.00001788151,0.0001558402,0.0009797101,0.00045406717,0.0023991058],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99752194,0.00048625923,0.00027886155,0.00060608867,0.0009153953,0.0001914839],"domain_scores_gemma":[0.98758686,0.0058223903,0.0016493861,0.0029598116,0.0018316797,0.00014994753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020500869,0.0009791012,0.0009309229,0.0033718885,0.0005906879,0.0012464611,0.001281695,0.0012943402,0.0022603483],"category_scores_gemma":[0.011954004,0.00048398238,0.0009044524,0.0016024894,0.0007039117,0.0017883537,0.0012358662,0.0010951713,0.0008587104],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035774402,0.00028762268,0.02896779,0.0007730835,0.00012834229,0.0010316641,0.00067926355,0.019572046,0.123826794,0.012202631,0.006026841,0.8061461],"study_design_scores_gemma":[0.00014244813,0.0006147914,0.02228917,0.00027603967,0.00023221837,0.0052555185,0.00046673344,0.59781814,0.2796851,0.04799639,0.044995002,0.00022843338],"about_ca_topic_score_codex":0.0015629259,"about_ca_topic_score_gemma":0.003078036,"teacher_disagreement_score":0.0033718885,"about_ca_system_score_codex":0.00043845226,"about_ca_system_score_gemma":0.0010284072,"threshold_uncertainty_score":0.010842025},"labels":[],"label_agreement":null},{"id":"W2009864036","doi":"10.1109/mtd.2014.10","title":"When-to-Release Decisions in Consideration of Technical Debt","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Technical debt; Context (archaeology); Debt; Risk analysis (engineering); Computer science; Duration (music); Software; Product (mathematics); Process (computing); Position (finance); Process management; Iterative and incremental development; New product development; Business; Software development; Software engineering; Finance; Marketing","score_opus":0.021377452541898614,"score_gpt":0.28926272299428607,"score_spread":0.26788527045238747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2009864036","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.68004954,0.0025043688,0.26029932,0.0072039166,0.0002625536,0.0012441863,0.00053212605,0.00037272464,0.04753122],"genre_scores_gemma":[0.9500965,0.00052995404,0.047107585,0.00021268881,0.00005320193,0.00014576348,0.00015077136,0.00007153819,0.0016319436],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9910188,0.0050159246,0.00064072927,0.00080730644,0.0014692086,0.0010480477],"domain_scores_gemma":[0.95680904,0.032940265,0.004755711,0.00088431133,0.0028054097,0.001805204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015558364,0.0012420577,0.0012068948,0.002358223,0.0016906841,0.0064725233,0.0017455558,0.002902171,0.0052213343],"category_scores_gemma":[0.043186497,0.0011972643,0.0009134846,0.00202993,0.0014657847,0.0057032434,0.002107275,0.0029427845,0.00036351677],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092920737,0.00066397415,0.031761304,0.00069944235,0.00035723168,0.004100866,0.0020572913,0.7530839,0.0066742683,0.05411319,0.004792983,0.1407664],"study_design_scores_gemma":[0.00010995071,0.00091776234,0.020604055,0.00034094063,0.00030007196,0.00073738967,0.005015261,0.89381474,0.0034300766,0.06632823,0.008150453,0.00025108564],"about_ca_topic_score_codex":0.0049835606,"about_ca_topic_score_gemma":0.0079363305,"teacher_disagreement_score":0.015558364,"about_ca_system_score_codex":0.0029057306,"about_ca_system_score_gemma":0.004195218,"threshold_uncertainty_score":0.08228147},"labels":[],"label_agreement":null},{"id":"W2010304830","doi":"10.1142/s0218488508005650","title":"CALIBRATING FUNCTION POINT BACKFIRING CONVERSION RATIOS USING NEURO-FUZZY TECHNIQUE","year":2008,"lang":"en","type":"article","venue":"International Journal of Uncertainty Fuzziness and Knowledge-Based Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Margin (machine learning); Fuzzy logic; Metric (unit); Computer science; Software; Calibration; Point (geometry); Artificial neural network; Sizing; Code (set theory); Algorithm; Machine learning; Mathematics; Artificial intelligence; Statistics; Engineering","score_opus":0.02986326649931817,"score_gpt":0.2729600884282638,"score_spread":0.24309682192894563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2010304830","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16186109,0.00013287111,0.8304767,0.00013632672,0.00004372257,0.00011045467,0.00006653659,0.0012353495,0.005936912],"genre_scores_gemma":[0.8556797,0.00007934191,0.14322864,0.000028838851,0.00000560647,0.000059935322,0.00006515629,0.000053034564,0.0007996805],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99841917,0.00042252685,0.00009119189,0.00030448768,0.0006819251,0.00008079754],"domain_scores_gemma":[0.99698466,0.001084477,0.00037387043,0.0004137293,0.0011143107,0.000028865887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025412792,0.0008045155,0.00058417226,0.0018612072,0.0005241046,0.0015379793,0.0011394466,0.0010146734,0.0009192422],"category_scores_gemma":[0.011889686,0.0004607741,0.0004940964,0.0010859382,0.0004588123,0.0013147277,0.00062226166,0.0008695546,0.00037677414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017813807,0.00014367542,0.007560709,0.00010743413,0.00006749285,0.00013949879,0.00044282127,0.77161646,0.012872082,0.0053808303,0.0009865905,0.20050427],"study_design_scores_gemma":[0.000011935918,0.000046990794,0.0018373032,0.000020796155,0.000017672373,0.000036543242,0.0000729193,0.98536646,0.010168081,0.0017532088,0.00064146955,0.000026572672],"about_ca_topic_score_codex":0.01043354,"about_ca_topic_score_gemma":0.007240646,"teacher_disagreement_score":0.01043354,"about_ca_system_score_codex":0.0013087802,"about_ca_system_score_gemma":0.0010405708,"threshold_uncertainty_score":0.020745635},"labels":[],"label_agreement":null},{"id":"W2010312922","doi":"10.1145/2245276.2231970","title":"IDE-based real-time focused search for near-miss clones","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Code (set theory); Software maintenance; Programming language; Software; Operating system; Software development; Software engineering; Biology; Gene; Genetics; Set (abstract data type)","score_opus":0.0323594242992459,"score_gpt":0.2966529840571515,"score_spread":0.26429355975790564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2010312922","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12936196,0.0010895396,0.7779053,0.00025182767,0.00028563093,0.00029932667,0.0011378459,0.08345581,0.006212784],"genre_scores_gemma":[0.34377486,0.00020730148,0.6450634,0.00026829902,0.00008206625,0.00027169526,0.0024418433,0.0021820285,0.005708606],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975032,0.00044276533,0.00019099518,0.0006598955,0.0010330459,0.00017008132],"domain_scores_gemma":[0.9845503,0.008110375,0.0012854635,0.001657619,0.0036334142,0.0007627567],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002723361,0.0011927998,0.0015030734,0.003649581,0.00044379532,0.0015093646,0.00213924,0.0016022079,0.004364721],"category_scores_gemma":[0.01245513,0.00059437036,0.0004973069,0.001124803,0.00033957377,0.0021344502,0.0015704675,0.0010557943,0.0032580765],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033051185,0.00057456805,0.019276839,0.0006149806,0.00018359149,0.0013327365,0.00060573034,0.009031185,0.1694949,0.0023226417,0.020578612,0.7726791],"study_design_scores_gemma":[0.0006180257,0.0012263348,0.0122679,0.000103234896,0.0002004661,0.0022647765,0.00033814315,0.77756476,0.17789711,0.0042055454,0.023117354,0.00019632286],"about_ca_topic_score_codex":0.00065625773,"about_ca_topic_score_gemma":0.001157306,"teacher_disagreement_score":0.004364721,"about_ca_system_score_codex":0.00034001062,"about_ca_system_score_gemma":0.00089929905,"threshold_uncertainty_score":0.014601469},"labels":[],"label_agreement":null},{"id":"W2010754032","doi":"10.1109/vissof.2007.4290711","title":"YARN: Animating Software Evolution","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Animation; Yarn; Software; Software evolution; Computer graphics (images); Computer animation; Reverse engineering; Software system; Code (set theory); Engineering drawing; Programming language; Software construction; Engineering","score_opus":0.014228935246945075,"score_gpt":0.26827113598879104,"score_spread":0.25404220074184597,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2010754032","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020428218,0.00028752466,0.96010494,0.0005642526,0.00011974454,0.00011814503,0.0003419357,0.011629534,0.006405645],"genre_scores_gemma":[0.12948112,0.0005685505,0.86089003,0.0002489711,0.000047072568,0.0003133144,0.00058095576,0.0024734729,0.005396449],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99957377,0.00019262068,0.00002147834,0.000074959906,0.00011289458,0.000024314204],"domain_scores_gemma":[0.9979594,0.0014827696,0.00009761364,0.00026165202,0.00012600153,0.00007256874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012554293,0.000718577,0.00040766055,0.0007561986,0.0004938441,0.0013114873,0.0015706111,0.0013909698,0.0076540727],"category_scores_gemma":[0.0045306482,0.0005742178,0.00076621026,0.00045332266,0.0010144071,0.0018981934,0.0018470075,0.0015491077,0.0010188726],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054392277,0.00024724222,0.0047949087,0.0014359104,0.00018053419,0.00093680044,0.009272906,0.21767455,0.103098266,0.13165,0.03411654,0.4960484],"study_design_scores_gemma":[0.00018660752,0.00024026848,0.0015633816,0.00024816982,0.00008395064,0.00063114177,0.00058999856,0.7079008,0.032051552,0.057676055,0.19871582,0.00011215949],"about_ca_topic_score_codex":0.0013323418,"about_ca_topic_score_gemma":0.0018458316,"teacher_disagreement_score":0.0076540727,"about_ca_system_score_codex":0.00042360282,"about_ca_system_score_gemma":0.00036464236,"threshold_uncertainty_score":0.02560544},"labels":[],"label_agreement":null},{"id":"W2010981136","doi":"10.1007/s11334-005-0011-3","title":"Tracing requirements to defect reports: an application of information retrieval techniques","year":2005,"lang":"en","type":"article","venue":"Innovations in Systems and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University; National Aeronautics and Space Administration","keywords":"Tracing; Computer science; TRACE (psycholinguistics); Debugging; Information retrieval; Precision and recall; Relevance (law); Software; Data mining; Hierarchy; Software engineering; Programming language","score_opus":0.014026146464856198,"score_gpt":0.2750939962476525,"score_spread":0.2610678497827963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2010981136","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04215214,0.0012699482,0.94437444,0.00029469133,0.00006233734,0.00041970587,0.00066398625,0.0067304187,0.004032348],"genre_scores_gemma":[0.15293337,0.0011397385,0.8408993,0.00008988922,0.00009629598,0.00019235454,0.0014371409,0.00036158052,0.0028503137],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.995516,0.0010562134,0.0004540201,0.00053441565,0.0022607867,0.00017862485],"domain_scores_gemma":[0.9806958,0.011942972,0.0017442994,0.0024522159,0.0029838432,0.00018093194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033994236,0.0013821266,0.001938134,0.0180833,0.0010878174,0.003294938,0.0024261903,0.001871497,0.003079223],"category_scores_gemma":[0.022361303,0.00057836785,0.0013675713,0.009589844,0.00067294517,0.003871656,0.0015157206,0.00093237165,0.0022455084],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002816953,0.0004387565,0.0050681005,0.0006955503,0.00016145143,0.0004320603,0.0012065275,0.0108412765,0.018908672,0.0069905743,0.0043436335,0.9506317],"study_design_scores_gemma":[0.0003088881,0.0010030539,0.01726681,0.00034989123,0.0011500464,0.0042776945,0.0018000127,0.7881712,0.10704107,0.04442122,0.033842802,0.0003672414],"about_ca_topic_score_codex":0.0056411624,"about_ca_topic_score_gemma":0.0039434135,"teacher_disagreement_score":0.0180833,"about_ca_system_score_codex":0.0008607219,"about_ca_system_score_gemma":0.0017759155,"threshold_uncertainty_score":0.017978072},"labels":[],"label_agreement":null},{"id":"W2011037576","doi":"10.5555/2486788.2487019","title":"Informing development decisions: from data to information","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software analytics; Software development; Software development process; Software engineering; Personal software process; Software project management; Software peer review; Data science; Package development process; Context (archaeology); Software quality; Software; Software construction; Knowledge management","score_opus":0.08094301475176306,"score_gpt":0.3005112809245753,"score_spread":0.21956826617281222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2011037576","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22636275,0.024698166,0.5757417,0.067938246,0.00091024773,0.0018261869,0.05379343,0.005032209,0.04369708],"genre_scores_gemma":[0.6124527,0.013529792,0.3346178,0.0026464462,0.00073174585,0.0007916777,0.033117093,0.00041012635,0.0017025417],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98009986,0.008899921,0.0023893616,0.002410806,0.005699497,0.0005006074],"domain_scores_gemma":[0.8201888,0.135079,0.011782659,0.018676788,0.012150897,0.0021218667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01803161,0.0017452561,0.0014776187,0.018904088,0.0012051597,0.015876004,0.002518301,0.003718623,0.002060802],"category_scores_gemma":[0.10932497,0.0015747577,0.00082749187,0.018463098,0.0029654938,0.025430627,0.0058960607,0.0041480553,0.0015807984],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061496726,0.0011617439,0.19403996,0.004505146,0.00062767445,0.0013593539,0.010496232,0.016703777,0.0061197155,0.057129648,0.031086935,0.6761549],"study_design_scores_gemma":[0.00018571907,0.0005005676,0.098141395,0.006903358,0.0007129069,0.0014406458,0.027652752,0.12021605,0.016642805,0.53508526,0.19188608,0.00063248957],"about_ca_topic_score_codex":0.004327802,"about_ca_topic_score_gemma":0.0046072123,"teacher_disagreement_score":0.018904088,"about_ca_system_score_codex":0.0021259766,"about_ca_system_score_gemma":0.00425872,"threshold_uncertainty_score":0.09536141},"labels":[],"label_agreement":null},{"id":"W2011315508","doi":"10.1109/icsm.2011.6080836","title":"Webdiff: A generic differencing service for software artifacts","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artifact (error); Software engineering; Software construction; Unified Modeling Language; Software development; Programming language; Software evolution; Software system; Software maintenance; Software; Source code; Artificial intelligence","score_opus":0.08138754593277489,"score_gpt":0.25570101701616604,"score_spread":0.17431347108339115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2011315508","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008596436,0.0004962214,0.7725771,0.00035551182,0.00020128471,0.00048198854,0.0039702486,0.20808393,0.0052373703],"genre_scores_gemma":[0.15764602,0.0013164574,0.73952335,0.0013613093,0.00034143036,0.0017164898,0.034534823,0.042358715,0.02120137],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9948704,0.0005994054,0.00084552757,0.00056367804,0.002824878,0.00029618182],"domain_scores_gemma":[0.9856792,0.00408369,0.0011484023,0.006610604,0.0018854789,0.00059269025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052995845,0.0019030762,0.0023193394,0.006879114,0.0010143559,0.00333487,0.00399663,0.0033119486,0.010829906],"category_scores_gemma":[0.02051942,0.001180941,0.0021346067,0.0050290725,0.0009105899,0.0062515857,0.009028899,0.0023039938,0.008142882],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018176808,0.00047523633,0.0132947285,0.0018289695,0.00040325857,0.0031973906,0.002195406,0.006795628,0.030620007,0.04214666,0.102197476,0.7950276],"study_design_scores_gemma":[0.00078690395,0.0002970665,0.009380937,0.0005376382,0.00021572017,0.0046001347,0.0007378815,0.21381605,0.06292744,0.097468816,0.6086479,0.00058348826],"about_ca_topic_score_codex":0.0029231121,"about_ca_topic_score_gemma":0.0027252962,"teacher_disagreement_score":0.010829906,"about_ca_system_score_codex":0.0014437924,"about_ca_system_score_gemma":0.0014351235,"threshold_uncertainty_score":0.03622961},"labels":[],"label_agreement":null},{"id":"W201142559","doi":"10.1007/978-3-642-38977-1_5","title":"An Assessment of Test-Driven Reuse: Promises and Pitfalls","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Computer science; Software engineering; Code reuse; Source code; Test (biology); Code (set theory); Work (physics); Test case; Programming language; Data science; Software; Machine learning; Engineering","score_opus":0.017282294242032262,"score_gpt":0.2994893514469103,"score_spread":0.28220705720487804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W201142559","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47323212,0.01586932,0.41848704,0.01044925,0.0003492152,0.00050991285,0.00090882735,0.0037018938,0.07649243],"genre_scores_gemma":[0.9264423,0.0017719491,0.06713116,0.00046963667,0.000076148346,0.00011551837,0.000383702,0.00032249917,0.0032871643],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9664302,0.012459998,0.0011770052,0.0010369373,0.018160658,0.00073528185],"domain_scores_gemma":[0.7652925,0.16454422,0.009623151,0.025944509,0.03224273,0.0023528386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030045968,0.0009918049,0.0008767905,0.004656135,0.0008014308,0.004145222,0.0037029264,0.0020133685,0.003726379],"category_scores_gemma":[0.12524264,0.00037840588,0.0011317687,0.0040223156,0.0025434534,0.006490815,0.0028089187,0.0016201431,0.0009891372],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008441078,0.0006690456,0.039473847,0.00097326527,0.0002979139,0.00017357554,0.0015428878,0.022145612,0.004846347,0.060414664,0.0053499076,0.86326885],"study_design_scores_gemma":[0.0004316826,0.009893026,0.1141018,0.0023145145,0.0011274037,0.0033062657,0.0045042336,0.40608704,0.04851285,0.35613924,0.05311005,0.00047195636],"about_ca_topic_score_codex":0.003722416,"about_ca_topic_score_gemma":0.0040797694,"teacher_disagreement_score":0.030045968,"about_ca_system_score_codex":0.0021497163,"about_ca_system_score_gemma":0.0036083711,"threshold_uncertainty_score":0.1589002},"labels":[],"label_agreement":null},{"id":"W2011458734","doi":"10.1145/1449764.1449790","title":"Enabling static analysis for partial java programs","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":143,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Programming language; Source code; Java; Static analysis; Declaration; Static program analysis; Class (philosophy); Theoretical computer science; Software; Software development; Artificial intelligence","score_opus":0.05940412544007045,"score_gpt":0.3003256130882066,"score_spread":0.24092148764813615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2011458734","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057699554,0.00030198423,0.9105097,0.00031506454,0.00006364078,0.00008431234,0.00048032327,0.024512498,0.0060329265],"genre_scores_gemma":[0.59929705,0.00089305674,0.38710243,0.00026007704,0.00018497604,0.0002875181,0.0025209633,0.005330321,0.0041235373],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970433,0.0006486296,0.00022761118,0.0003955212,0.0012318874,0.00045307478],"domain_scores_gemma":[0.9892313,0.006797294,0.00055264984,0.0020873402,0.0011418716,0.00018958644],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020427397,0.0010197146,0.0012159003,0.0024686933,0.0011344115,0.0031342548,0.0014370087,0.0008451049,0.004435593],"category_scores_gemma":[0.014384464,0.0011673694,0.0014297339,0.0020500205,0.0021134042,0.005339552,0.00470062,0.0014547561,0.0016956241],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011703005,0.0002508499,0.011013112,0.0015316507,0.00016887455,0.0014438623,0.0027100858,0.08938359,0.13299194,0.24599186,0.013849616,0.49949428],"study_design_scores_gemma":[0.00010359753,0.00013222932,0.0025300863,0.00029521703,0.00016165191,0.0005149796,0.00039281027,0.55351174,0.11213915,0.29715633,0.03294771,0.00011447337],"about_ca_topic_score_codex":0.0027040308,"about_ca_topic_score_gemma":0.004749798,"teacher_disagreement_score":0.004435593,"about_ca_system_score_codex":0.0010472115,"about_ca_system_score_gemma":0.0018321964,"threshold_uncertainty_score":0.014838576},"labels":[],"label_agreement":null},{"id":"W2011861933","doi":"10.1109/csmr.2012.78","title":"A Comparative Study of the Performance of IR Models on Duplicate Bug Detection","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Heuristics; Weighting; Set (abstract data type); Data mining; Entropy (arrow of time); Machine learning; Artificial intelligence; Information retrieval; Programming language","score_opus":0.057317529597373316,"score_gpt":0.29683549624866024,"score_spread":0.2395179666512869,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2011861933","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8403862,0.016668739,0.11355882,0.0019023587,0.0006633191,0.0005405686,0.001629085,0.014820646,0.009830367],"genre_scores_gemma":[0.92873055,0.0027376015,0.062078986,0.00020018627,0.00023695205,0.00020577703,0.0026038282,0.00051431824,0.0026919325],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98611337,0.008077946,0.0014203178,0.0017310036,0.0020539989,0.00060333865],"domain_scores_gemma":[0.90748036,0.07665119,0.0023261455,0.0055433647,0.007043684,0.00095510075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030875083,0.002709652,0.0027658343,0.0053758565,0.0008882153,0.0028628851,0.0019158316,0.0024869482,0.00099027],"category_scores_gemma":[0.06101016,0.0008819447,0.0018695927,0.0036016142,0.00085725763,0.0064460463,0.0014924443,0.0017441688,0.0015693974],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0056519927,0.0019565623,0.036414266,0.0024199407,0.0026736187,0.00029699082,0.0012439644,0.30166796,0.010460356,0.0021893284,0.013059949,0.6219651],"study_design_scores_gemma":[0.00012133349,0.0019846873,0.009084955,0.00009949127,0.0005038677,0.00026893045,0.0004192192,0.975878,0.008132055,0.0015265843,0.0018282194,0.00015275834],"about_ca_topic_score_codex":0.0077703204,"about_ca_topic_score_gemma":0.005413928,"teacher_disagreement_score":0.030875083,"about_ca_system_score_codex":0.001770201,"about_ca_system_score_gemma":0.0012530772,"threshold_uncertainty_score":0.16328502},"labels":[],"label_agreement":null},{"id":"W2012665504","doi":"10.1145/1923947.1923954","title":"F007","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; TRACE (psycholinguistics); Suite; Code (set theory); Function (biology); Identification (biology); Field (mathematics); Source code; Fault (geology); Programming language; Mathematics; Geography; Geology","score_opus":0.010596219187848527,"score_gpt":0.2607921228247907,"score_spread":0.25019590363694216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2012665504","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010234289,0.0014584645,0.04047072,0.004418939,0.003007213,0.00060190103,0.020512052,0.027673861,0.8916225],"genre_scores_gemma":[0.053158298,0.0019254044,0.04001989,0.0038837232,0.0010323152,0.00043976272,0.046108328,0.005772546,0.84765977],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99828404,0.00020142151,0.00009932005,0.00035316273,0.00084365625,0.00021830012],"domain_scores_gemma":[0.995654,0.0009559361,0.00025243487,0.0006601743,0.0020263789,0.00045096976],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0015124435,0.001127986,0.00062345096,0.0035014711,0.0019291956,0.0045414185,0.0017462615,0.0025850344,0.48915133],"category_scores_gemma":[0.0077729817,0.0005405709,0.0006879151,0.0024878487,0.00064351474,0.0028082244,0.002274074,0.0012429968,0.41185692],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019934025,0.000090597416,0.003200169,0.00033181437,0.00001637157,0.000304509,0.00019951601,0.00049011776,0.002379805,0.009269694,0.5450118,0.43850622],"study_design_scores_gemma":[0.00002099551,0.00005901036,0.002044961,0.0000941221,0.000008374399,0.00041583405,0.00011871954,0.0010952973,0.001440847,0.0026594235,0.9920204,0.000022027734],"about_ca_topic_score_codex":0.0056641717,"about_ca_topic_score_gemma":0.0066629197,"teacher_disagreement_score":0.51084864,"about_ca_system_score_codex":0.0014789735,"about_ca_system_score_gemma":0.0022370163,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2012991702","doi":"10.1287/isre.1050.0079","title":"Conceptualizing Systems for Understanding: An Empirical Test of Decomposition Principles in Object-Oriented Analysis","year":2006,"lang":"en","type":"article","venue":"Information Systems Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":162,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"University of British Columbia; University of Georgia","keywords":"Operationalization; Domain model; Computer science; Unified Modeling Language; Domain (mathematical analysis); Meaning (existential); Empirical research; Cohesion (chemistry); Conceptual model; Domain analysis; Knowledge management; Psychology; Epistemology; Domain knowledge; Software development; Programming language; Mathematics","score_opus":0.14899362967571214,"score_gpt":0.40959886085126274,"score_spread":0.26060523117555057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2012991702","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99199206,0.000057832967,0.004829325,0.00023589809,0.000011944652,0.00033591042,0.000019468906,0.00001747723,0.002500087],"genre_scores_gemma":[0.99297875,0.000035630383,0.0061345967,0.0000998772,0.0000078676685,0.00049269537,0.000044270306,0.000017312941,0.00018900006],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.96401215,0.024319615,0.002158997,0.0035649403,0.0049645947,0.000979682],"domain_scores_gemma":[0.57941693,0.3642319,0.022201715,0.019782793,0.011468341,0.0028983485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06503344,0.0005940645,0.0005929419,0.0015985948,0.0017019097,0.0039206347,0.0019840628,0.0017400658,0.0034301614],"category_scores_gemma":[0.27215737,0.00072578207,0.0007776157,0.00091462024,0.008111228,0.012679977,0.0063522737,0.0043141674,0.00033243053],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004956113,0.010291731,0.3619522,0.0016724477,0.00052038435,0.0004674336,0.41589624,0.0046043256,0.00994781,0.042028125,0.0019198317,0.14574331],"study_design_scores_gemma":[0.0027396635,0.017707864,0.496324,0.0011989846,0.00062090944,0.0012617237,0.25858048,0.09950453,0.01158494,0.09173584,0.018342594,0.00039852477],"about_ca_topic_score_codex":0.0008655788,"about_ca_topic_score_gemma":0.00052759005,"teacher_disagreement_score":0.06503344,"about_ca_system_score_codex":0.0014165343,"about_ca_system_score_gemma":0.0015423512,"threshold_uncertainty_score":0.34393382},"labels":[],"label_agreement":null},{"id":"W2013410000","doi":"10.1145/1353482.1353501","title":"View-based maintenance of graphical user interfaces","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Graphical user interface; Source code; Software maintenance; Context (archaeology); Graphical user interface testing; Programming language; Human–computer interaction; User interface; Software engineering; Object (grammar); Code (set theory); Plug-in; Object-oriented programming; Software; Interface (matter); Software system; User interface design; Operating system; Set (abstract data type); Artificial intelligence","score_opus":0.020872006740764087,"score_gpt":0.25512254375783827,"score_spread":0.23425053701707418,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2013410000","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05088458,0.00058295747,0.9279976,0.00028757384,0.000065800225,0.0000974768,0.0001593901,0.015985414,0.0039392523],"genre_scores_gemma":[0.57330173,0.00046425103,0.41534728,0.000299042,0.000086823085,0.00022056831,0.0012856169,0.0045126313,0.0044821543],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99560994,0.0009611211,0.00030614925,0.00090610416,0.0019598627,0.00025693903],"domain_scores_gemma":[0.97026044,0.0069939634,0.0021646977,0.015915371,0.0043345205,0.00033102915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046651405,0.00080788846,0.000776067,0.0017050672,0.0006601067,0.0026332128,0.0033211366,0.00134101,0.0013943168],"category_scores_gemma":[0.02531141,0.001179142,0.00093646126,0.00092585455,0.0012891446,0.0046697105,0.0031092123,0.0018560084,0.00072161434],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040464435,0.00030547607,0.02009502,0.0006505202,0.00020483899,0.0007860366,0.0032529922,0.042953588,0.0909644,0.044661254,0.009256017,0.7864652],"study_design_scores_gemma":[0.00019953109,0.0005951322,0.014159012,0.00033003956,0.0006928507,0.0024656232,0.0006289285,0.6197956,0.22148626,0.07748368,0.061906952,0.00025642337],"about_ca_topic_score_codex":0.0017036881,"about_ca_topic_score_gemma":0.0018553832,"teacher_disagreement_score":0.0046651405,"about_ca_system_score_codex":0.0009134145,"about_ca_system_score_gemma":0.0011672905,"threshold_uncertainty_score":0.024671912},"labels":[],"label_agreement":null},{"id":"W2013839424","doi":"10.1145/1137983.1138030","title":"Examining the evolution of code comments in PostgreSQL","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":81,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code (set theory); Computer science; Software evolution; Constant (computer programming); Software; Code review; Software engineering; Programming language; Style (visual arts); Source code; Software development; Database; Software quality; Software construction; History","score_opus":0.02549286550906637,"score_gpt":0.263383876298708,"score_spread":0.23789101078964164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2013839424","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.995621,0.000110579436,0.0014080242,0.00013119729,0.00001227021,0.000036873545,0.0015516378,0.00047241658,0.00065593555],"genre_scores_gemma":[0.99004173,0.00005690755,0.0031776389,0.000050580842,0.000022989394,0.00006627695,0.0047128326,0.00023254631,0.0016383653],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99230677,0.0016706062,0.00052681234,0.0011453764,0.003929272,0.00042118848],"domain_scores_gemma":[0.8212227,0.084040314,0.03760444,0.010329217,0.043813996,0.0029893056],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057031135,0.00036972424,0.0003147868,0.0043490254,0.0005688895,0.0011308031,0.0007973061,0.000850466,0.0007038527],"category_scores_gemma":[0.06748386,0.00037760875,0.00028394806,0.0042340797,0.0007991886,0.0013301838,0.00084496,0.0011118832,0.0006007674],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008462488,0.0006406117,0.83643967,0.00043499755,0.00015174678,0.0009983993,0.017317895,0.005233631,0.023501161,0.0005705597,0.0060360236,0.10782906],"study_design_scores_gemma":[0.00001738338,0.00028984182,0.97720987,0.00003011641,0.000018535755,0.0002749598,0.0013417003,0.010987876,0.0048345965,0.00012053667,0.004818038,0.000056644432],"about_ca_topic_score_codex":0.015842455,"about_ca_topic_score_gemma":0.022009226,"teacher_disagreement_score":0.015842455,"about_ca_system_score_codex":0.0013722158,"about_ca_system_score_gemma":0.00083832716,"threshold_uncertainty_score":0.03150052},"labels":[],"label_agreement":null},{"id":"W2014047795","doi":"10.1145/1297846.1297960","title":"Detection and correction of design defects in object-oriented designs","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Object-oriented design; Object (grammar); Error detection and correction; Quality (philosophy); Object-oriented programming; Algorithm; Computer engineering; Programming language; Artificial intelligence","score_opus":0.025787682790200286,"score_gpt":0.2662681915689593,"score_spread":0.24048050877875904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2014047795","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03805993,0.00022973614,0.95481944,0.00024742755,0.000039249062,0.00024341064,0.00009855551,0.0055367392,0.000725491],"genre_scores_gemma":[0.1779142,0.00019519466,0.819746,0.00011611681,0.00001714699,0.00021340536,0.00031295797,0.000691238,0.0007936567],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98321044,0.0054681534,0.0019242413,0.0018641342,0.0070365076,0.00049658125],"domain_scores_gemma":[0.93984056,0.02871313,0.010596812,0.01274885,0.007718327,0.0003823194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071064476,0.0015288783,0.0011190611,0.0033002729,0.00078829104,0.0021053345,0.0023567348,0.0024961673,0.00075755746],"category_scores_gemma":[0.054454204,0.0012963025,0.0013843108,0.0011331959,0.001961472,0.0025908882,0.0018027846,0.0019717258,0.00046689104],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027512314,0.0005243204,0.034288555,0.0016378225,0.00019591524,0.0013937983,0.0020606942,0.08444597,0.078607574,0.03746004,0.0035582045,0.755552],"study_design_scores_gemma":[0.0002201032,0.0007200067,0.011670375,0.0008987497,0.00037039607,0.0038374385,0.0006009714,0.55202365,0.3210621,0.08204709,0.026216643,0.0003325222],"about_ca_topic_score_codex":0.0013251025,"about_ca_topic_score_gemma":0.0016227422,"teacher_disagreement_score":0.0071064476,"about_ca_system_score_codex":0.0008834292,"about_ca_system_score_gemma":0.002772033,"threshold_uncertainty_score":0.037582994},"labels":[],"label_agreement":null},{"id":"W2014256464","doi":"10.1142/s0218194006002707","title":"UNDERSTANDING THE EVOLUTION AND CO-EVOLUTION OF CLASSES IN OBJECT-ORIENTED SYSTEMS","year":2006,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Class diagram; Class (philosophy); Association rule learning; Software evolution; Unified Modeling Language; Software system; Data mining; Categorical variable; Software; Artificial intelligence; Programming language; Machine learning; Software construction","score_opus":0.016992741941098635,"score_gpt":0.24870536454836795,"score_spread":0.23171262260726933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2014256464","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7153863,0.001206458,0.27980962,0.0009667527,0.000019034918,0.00015398189,0.00036663134,0.000459581,0.0016316567],"genre_scores_gemma":[0.7852315,0.0006388094,0.21260531,0.00008225403,0.00002547292,0.00006866463,0.0007565433,0.00006154376,0.00052997784],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99830025,0.00045562958,0.00020342394,0.0003503663,0.0005894497,0.00010091934],"domain_scores_gemma":[0.98291093,0.010691402,0.0031149068,0.0017163755,0.0012952323,0.0002711777],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024424873,0.00040976724,0.00040891458,0.005857103,0.00086114474,0.0021482797,0.0012209556,0.0013100823,0.00032143656],"category_scores_gemma":[0.018097255,0.0005502606,0.00057118415,0.003572878,0.0012273518,0.004648332,0.001043305,0.0010744419,0.00013450938],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028506556,0.00080091006,0.39603573,0.0006452575,0.0001887427,0.0017297248,0.009447076,0.07473591,0.020044122,0.019386267,0.001260913,0.4754402],"study_design_scores_gemma":[0.00003982407,0.00022566303,0.2538128,0.0001966917,0.00017422729,0.0020925477,0.0038530242,0.6401536,0.015683696,0.07370557,0.009959941,0.000102345606],"about_ca_topic_score_codex":0.011321554,"about_ca_topic_score_gemma":0.014247598,"teacher_disagreement_score":0.011321554,"about_ca_system_score_codex":0.0010884062,"about_ca_system_score_gemma":0.00093602145,"threshold_uncertainty_score":0.022511303},"labels":[],"label_agreement":null},{"id":"W2015025073","doi":"10.1109/wse.2011.6081813","title":"Using indexed sequence diagrams to recover the behaviour of AJAX applications","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Ajax; Asynchronous communication; Computer science; Web application; World Wide Web; Computer network","score_opus":0.14775755673080543,"score_gpt":0.3310384812984912,"score_spread":0.1832809245676858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015025073","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08075411,0.00023712813,0.90779364,0.00016666617,0.000046631973,0.0002706245,0.00050569937,0.008165706,0.002059764],"genre_scores_gemma":[0.34802753,0.00044061904,0.64535207,0.000056980796,0.000030390884,0.00034026196,0.0019293596,0.001154171,0.0026686313],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966274,0.0010530453,0.00029867596,0.00046357437,0.0014029373,0.00015432561],"domain_scores_gemma":[0.9824257,0.010775417,0.0019005616,0.0025542115,0.0021327706,0.0002113772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031687582,0.00091418717,0.00038473305,0.0036676514,0.0006600073,0.0015073437,0.0010947097,0.0009635754,0.0019940273],"category_scores_gemma":[0.02246114,0.0006179971,0.000834126,0.0013757979,0.0010491413,0.0024421108,0.0009909176,0.0011945947,0.00056239223],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075324724,0.00051164004,0.027683392,0.0013947254,0.00020901067,0.0029593287,0.007859842,0.35274687,0.084236406,0.12560683,0.003996094,0.39204258],"study_design_scores_gemma":[0.00007219958,0.00023194334,0.0048597227,0.00016130312,0.00011745689,0.00064385193,0.0004107577,0.885013,0.037735958,0.041752134,0.02891649,0.00008527529],"about_ca_topic_score_codex":0.005707744,"about_ca_topic_score_gemma":0.0054629473,"teacher_disagreement_score":0.005707744,"about_ca_system_score_codex":0.0012214308,"about_ca_system_score_gemma":0.0018124309,"threshold_uncertainty_score":0.016758204},"labels":[],"label_agreement":null},{"id":"W2015215117","doi":"10.1007/s11219-006-9218-2","title":"On evaluating the layout of UML diagrams for program comprehension","year":2006,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Class diagram; Sequence diagram; Applications of UML; Story-driven modeling; UML tool; Computer science; Unified Modeling Language; Communication diagram; Programming language; Activity diagram; Software engineering; Flowchart; Readability; Class (philosophy); Program comprehension; Use Case Diagram; Object Constraint Language; Software; Software system; Artificial intelligence","score_opus":0.08799721980182837,"score_gpt":0.4182070965023177,"score_spread":0.3302098767004893,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015215117","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6277957,0.0009175699,0.35761485,0.00040044263,0.00008516542,0.00044033947,0.00043100625,0.005814361,0.0065006036],"genre_scores_gemma":[0.7412964,0.0004133077,0.25390756,0.000080454134,0.00003669525,0.00014572807,0.0013574797,0.0007477857,0.0020145334],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99279803,0.004285528,0.0004497195,0.00064354716,0.0015801619,0.0002429421],"domain_scores_gemma":[0.8286266,0.15125868,0.0044054016,0.004115639,0.010062542,0.0015312122],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005629472,0.0012273496,0.00091106293,0.0035655198,0.00076099136,0.002717255,0.0012300086,0.0017397495,0.005346701],"category_scores_gemma":[0.10537454,0.0005318857,0.00061865203,0.001971996,0.00070100266,0.0051402305,0.0015345797,0.0008492525,0.0008472698],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025289415,0.0012624698,0.039576214,0.0010913343,0.00016311473,0.00023177818,0.004312283,0.044927735,0.07201586,0.0050852653,0.004218918,0.82458615],"study_design_scores_gemma":[0.0005493962,0.0044887797,0.065997966,0.00035116833,0.00046132438,0.00040598752,0.0047840076,0.82631606,0.0774894,0.012839931,0.006148449,0.0001675745],"about_ca_topic_score_codex":0.007655989,"about_ca_topic_score_gemma":0.011606581,"teacher_disagreement_score":0.007655989,"about_ca_system_score_codex":0.0012532478,"about_ca_system_score_gemma":0.0015966386,"threshold_uncertainty_score":0.029771805},"labels":[],"label_agreement":null},{"id":"W2015254062","doi":"10.5555/2819009.2819194","title":"Understanding the software fault introduction process","year":2015,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Debugging; Computer science; Software engineering; Process (computing); Software bug; Fault (geology); Software fault tolerance; Algorithmic program debugging; Meaning (existential); Software; Programming language","score_opus":0.12096392220138281,"score_gpt":0.313044506078589,"score_spread":0.19208058387720622,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015254062","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2738136,0.002305105,0.63724166,0.040823005,0.00022285759,0.0007983426,0.00024268836,0.00069490075,0.04385782],"genre_scores_gemma":[0.8486862,0.0021555712,0.14052095,0.0015643563,0.000080096295,0.00049163547,0.00035165064,0.00013760566,0.006011856],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9935708,0.004040292,0.00027297763,0.0006677019,0.0009797702,0.00046844236],"domain_scores_gemma":[0.9530723,0.033838343,0.0048946436,0.0030505708,0.004198766,0.0009453108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015789103,0.0010427545,0.00042018812,0.0036323348,0.002146074,0.008637749,0.0026737968,0.0070675686,0.004430567],"category_scores_gemma":[0.03917761,0.0010717241,0.0009875525,0.0017460404,0.013814055,0.023706716,0.0034152383,0.006508394,0.00080619275],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008208464,0.0002651398,0.014359078,0.00044215872,0.000020211643,0.0003717194,0.04613311,0.013724543,0.0017842309,0.8692013,0.0020924567,0.051524036],"study_design_scores_gemma":[0.00009071392,0.0002975212,0.01225117,0.00095071766,0.00007511747,0.0005232048,0.030245287,0.09098108,0.0046365024,0.8062581,0.05357604,0.00011456023],"about_ca_topic_score_codex":0.008377268,"about_ca_topic_score_gemma":0.0041191163,"teacher_disagreement_score":0.015789103,"about_ca_system_score_codex":0.0053055123,"about_ca_system_score_gemma":0.008710738,"threshold_uncertainty_score":0.083501756},"labels":[],"label_agreement":null},{"id":"W2015271973","doi":"10.1109/ms.2005.156","title":"Improving After-the-Fact Tracing and Mapping: Supporting Software Quality Predictions","year":2005,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University; National Aeronautics and Space Administration","keywords":"Tracing; Traceability; Computer science; Software engineering; Software; Quality (philosophy); Programming language","score_opus":0.023267198100036784,"score_gpt":0.2865996602850546,"score_spread":0.2633324621850178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015271973","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10834905,0.0002932362,0.8656391,0.00055830425,0.000037977756,0.0003061851,0.00067978,0.020447277,0.0036889962],"genre_scores_gemma":[0.43169528,0.00021401102,0.5646396,0.00006748853,0.000031972886,0.00016887047,0.001619412,0.00062503264,0.00093830837],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9910286,0.00368799,0.00072902726,0.0013667814,0.002834808,0.00035284497],"domain_scores_gemma":[0.9059756,0.06380642,0.0078007737,0.012611363,0.009060091,0.00074578176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011924371,0.0016078164,0.0010925912,0.007626878,0.0011218175,0.003949128,0.002230386,0.0018690359,0.0023130015],"category_scores_gemma":[0.12931715,0.00094023946,0.0007729105,0.0040042526,0.0010137719,0.007468712,0.0024861488,0.001440884,0.0010471791],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053564983,0.0007001945,0.041586094,0.000315007,0.0001167328,0.0003571765,0.0029817142,0.06809992,0.00835211,0.0137126595,0.005117632,0.8581251],"study_design_scores_gemma":[0.00007443935,0.00015152179,0.0058518676,0.000071342394,0.00006680493,0.0001458959,0.0003171965,0.9419175,0.018332176,0.02856897,0.0044420543,0.00006031203],"about_ca_topic_score_codex":0.008064673,"about_ca_topic_score_gemma":0.0066035897,"teacher_disagreement_score":0.011924371,"about_ca_system_score_codex":0.00096195587,"about_ca_system_score_gemma":0.002384263,"threshold_uncertainty_score":0.06306285},"labels":[],"label_agreement":null},{"id":"W2015538933","doi":"10.1109/csmr.2012.40","title":"Recommending Refactorings to Reverse Software Architecture Erosion","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Fundação de Amparo à Pesquisa do Estado de Minas Gerais; Japan Society for the Promotion of Science; Conselho Nacional de Desenvolvimento Científico e Tecnológico","keywords":"Code refactoring; Computer science; Software engineering; Architecture; Software architecture; Process (computing); Software; Programming language","score_opus":0.02731706643225493,"score_gpt":0.2810299039106431,"score_spread":0.25371283747838813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015538933","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27613187,0.0059643094,0.6859696,0.0031704886,0.00060839724,0.0021381206,0.0009967941,0.017061107,0.0079593435],"genre_scores_gemma":[0.30465338,0.003151881,0.68118316,0.00048883003,0.00026294187,0.00048781896,0.002407855,0.0008155427,0.0065485747],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9953288,0.0017725951,0.0004613552,0.0006295738,0.0015278219,0.00028002245],"domain_scores_gemma":[0.9740596,0.011060036,0.0020578671,0.00396601,0.008206347,0.0006502024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005826853,0.0020864694,0.0011128024,0.004524399,0.0010095182,0.0018470214,0.0024230175,0.0026804975,0.003375233],"category_scores_gemma":[0.03558963,0.0007501263,0.0011588406,0.0017951686,0.00039695628,0.0023007952,0.0008581321,0.0016661533,0.0020601656],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037494482,0.0011536266,0.021119205,0.0011567123,0.00032186147,0.00070825766,0.0014203782,0.020446995,0.03133189,0.0020397617,0.019128792,0.90079755],"study_design_scores_gemma":[0.0007378661,0.0028556518,0.035402093,0.00220354,0.002073469,0.0023102758,0.0047412077,0.68419373,0.11762703,0.01948608,0.12777844,0.00059055275],"about_ca_topic_score_codex":0.008180885,"about_ca_topic_score_gemma":0.028967824,"teacher_disagreement_score":0.008180885,"about_ca_system_score_codex":0.0007041053,"about_ca_system_score_gemma":0.0023657796,"threshold_uncertainty_score":0.030815661},"labels":[],"label_agreement":null},{"id":"W2015739841","doi":"10.1145/568760.568817","title":"Designing a component-based framework for visualization in software engineering and knowledge engineering","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"University of Victoria","keywords":"Computer science; Software engineering; Component (thermodynamics); Visualization; Domain (mathematical analysis); Business process reengineering; Component-based software engineering; Systems engineering; Software development; Software; Engineering; Data mining; Programming language","score_opus":0.029278816750086978,"score_gpt":0.27604245545388145,"score_spread":0.24676363870379447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015739841","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012722707,0.0003699568,0.9952441,0.0003664527,0.00003540822,0.00010605618,0.000027471322,0.001456447,0.0011217913],"genre_scores_gemma":[0.009515972,0.00029258247,0.9891395,0.000036632202,0.000014259064,0.00011858425,0.0000993834,0.00022162231,0.0005615462],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99699247,0.0014640498,0.00025364358,0.00036173,0.0007546546,0.00017343719],"domain_scores_gemma":[0.99644715,0.0016201047,0.00022709921,0.0006091777,0.0007581382,0.00033835007],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007826691,0.0013482409,0.00088231586,0.0041430434,0.0025149826,0.007190026,0.0029660733,0.0021630751,0.0032205766],"category_scores_gemma":[0.010826698,0.0013264612,0.0016490348,0.003528639,0.0026505801,0.008080759,0.0038477657,0.0032833433,0.0012359517],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000100538455,0.00016098208,0.0023021551,0.0013093962,0.00015387409,0.0008445652,0.012347758,0.03497443,0.018907655,0.41189402,0.017852861,0.49915183],"study_design_scores_gemma":[0.00013060281,0.00022587736,0.0015399547,0.0010310531,0.00018048265,0.0017774648,0.0025365565,0.18011242,0.013376616,0.30326748,0.49549785,0.00032363491],"about_ca_topic_score_codex":0.009200707,"about_ca_topic_score_gemma":0.01424396,"teacher_disagreement_score":0.009200707,"about_ca_system_score_codex":0.0017381005,"about_ca_system_score_gemma":0.0031247283,"threshold_uncertainty_score":0.04139197},"labels":[],"label_agreement":null},{"id":"W2015747019","doi":"10.1016/j.jss.2008.11.846","title":"Determining factors that affect long-term evolution in scientific application software","year":2008,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Software evolution; Software development; Social software engineering; Term (time); Software construction; Software; Software engineering; Computer science; Software analytics; Software sizing; Software system; Software maintenance; Personal software process; Operating system","score_opus":0.03252627983587554,"score_gpt":0.2681634616516386,"score_spread":0.2356371818157631,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015747019","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99876785,0.00010050202,0.0004171595,0.00005367837,0.0000034015904,0.000006440646,0.000023987206,0.000010120757,0.00061690994],"genre_scores_gemma":[0.99937063,0.000027922639,0.00026761275,0.000018867584,0.000004097449,0.0000034992004,0.00006180798,0.000013828801,0.00023183467],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9984843,0.0004529557,0.0001735457,0.00021817446,0.00046475136,0.00020628645],"domain_scores_gemma":[0.9605301,0.02391183,0.0074365283,0.0013030875,0.0052087074,0.0016096843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023506607,0.00019003532,0.00022212623,0.001016645,0.00061589817,0.0015747104,0.00040960964,0.00076432806,0.0017831296],"category_scores_gemma":[0.03371463,0.00022001924,0.00028019655,0.00086511875,0.0005234989,0.0011658986,0.00050589326,0.00088990905,0.00032079947],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042782797,0.0003682324,0.9414322,0.00005725975,0.00014859656,0.00024686428,0.0006034845,0.0020846385,0.03453902,0.0004216251,0.0003003964,0.019369893],"study_design_scores_gemma":[0.0000074579953,0.00019793578,0.9883167,0.00001010466,0.000068524714,0.00026234688,0.00050324196,0.004337799,0.0054053217,0.00048904505,0.00038467528,0.0000168998],"about_ca_topic_score_codex":0.0028208632,"about_ca_topic_score_gemma":0.005333197,"teacher_disagreement_score":0.0028208632,"about_ca_system_score_codex":0.00072262564,"about_ca_system_score_gemma":0.0005938615,"threshold_uncertainty_score":0.012431622},"labels":[],"label_agreement":null},{"id":"W2016035035","doi":"10.1109/csmr-wcre.2014.6747170","title":"Towards a context-aware IDE-based meta search engine for recommendation about programming errors and exceptions","year":2014,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Relevance (law); Search engine; Context (archaeology); World Wide Web; Information retrieval; Semantic search; Web search query; Web search engine; Popularity; Search analytics; Web crawler; Web page; Metasearch engine; Software; Programming language","score_opus":0.08063753346257617,"score_gpt":0.34394747021048877,"score_spread":0.2633099367479126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2016035035","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.292375,0.006689922,0.65899056,0.0009331492,0.00029543522,0.00085834006,0.004410601,0.026764955,0.008681991],"genre_scores_gemma":[0.4914786,0.0009133868,0.49799493,0.00039890758,0.00014965306,0.0002614587,0.005721086,0.00021641204,0.0028656477],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998979,0.00019964576,0.0001376918,0.00020704177,0.00038104277,0.00009564339],"domain_scores_gemma":[0.9977717,0.00076635904,0.00022207871,0.00028098238,0.0007684279,0.00019045355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014216041,0.00089497457,0.0017141472,0.006657503,0.0005646736,0.0014882088,0.0013668513,0.0013537452,0.00097017974],"category_scores_gemma":[0.0042199036,0.00043688877,0.0010329984,0.0029395504,0.00015743056,0.0019947696,0.000995939,0.00094917003,0.0011472671],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016983497,0.0015846039,0.06815178,0.0010413515,0.0006287735,0.0010871212,0.0006135039,0.021361288,0.043885697,0.004528602,0.020874089,0.8345449],"study_design_scores_gemma":[0.00025894385,0.00052816985,0.018303832,0.00012182457,0.00042870303,0.0008901191,0.0004328009,0.94697374,0.017837266,0.0039003564,0.010211935,0.000112309695],"about_ca_topic_score_codex":0.009187875,"about_ca_topic_score_gemma":0.023243003,"teacher_disagreement_score":0.009187875,"about_ca_system_score_codex":0.0004518373,"about_ca_system_score_gemma":0.0013915024,"threshold_uncertainty_score":0.018268824},"labels":[],"label_agreement":null},{"id":"W2016665229","doi":"10.5555/2662708.2662712","title":"Refactoring clones: a new perspective","year":2013,"lang":"en","type":"article","venue":"International Workshop on Software Clones","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; clone (Java method); Computer science; Perspective (graphical); Software engineering; Software evolution; Software; Position paper; Programming language; Software development; Artificial intelligence; World Wide Web; Software construction","score_opus":0.028413841104992916,"score_gpt":0.30573464085672647,"score_spread":0.27732079975173357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2016665229","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010917683,0.1546045,0.61376435,0.14643477,0.0046005263,0.00008625352,0.00012106827,0.0006720531,0.06879881],"genre_scores_gemma":[0.3371057,0.2118305,0.38944355,0.020127414,0.01568043,0.00032745348,0.0002419235,0.0007758041,0.02446721],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98592585,0.006589143,0.0008477067,0.0019706858,0.0041396315,0.0005269894],"domain_scores_gemma":[0.9444729,0.041272562,0.002023581,0.0049701966,0.005871368,0.001389418],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017076679,0.0019375697,0.0026292447,0.007938069,0.0034809539,0.0144028105,0.0044528376,0.009674592,0.0046562348],"category_scores_gemma":[0.021931432,0.0011981871,0.0018240276,0.0062505226,0.025226707,0.045125615,0.006184084,0.012358208,0.0012464729],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030739717,0.000043447293,0.00033101067,0.0005233065,0.000021522621,0.0002624889,0.0030422918,0.00070221897,0.0006096029,0.9436985,0.0033195477,0.047415305],"study_design_scores_gemma":[0.000037405847,0.00013008906,0.0002946494,0.00069529796,0.0000478117,0.0008662519,0.002796493,0.0044612098,0.0009272837,0.82825327,0.16142581,0.00006438641],"about_ca_topic_score_codex":0.002687897,"about_ca_topic_score_gemma":0.0023757752,"teacher_disagreement_score":0.017076679,"about_ca_system_score_codex":0.0047950647,"about_ca_system_score_gemma":0.003724788,"threshold_uncertainty_score":0.09031123},"labels":[],"label_agreement":null},{"id":"W2016902117","doi":"10.1109/icsm.2013.47","title":"Refactoring Clones: An Optimization Problem","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Reusability; Computer science; Matching (statistics); Process (computing); Software; Identifier; Programming language; Software maintenance; Software system; Mathematics","score_opus":0.019255559285741995,"score_gpt":0.2554330205148758,"score_spread":0.23617746122913383,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2016902117","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06352489,0.0007133224,0.9294559,0.0014093759,0.000063741376,0.0002980968,0.00026909134,0.00070509774,0.003560473],"genre_scores_gemma":[0.17975545,0.00040590126,0.81271285,0.00032803213,0.00007209167,0.0003808817,0.0005986224,0.00043153565,0.0053146933],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968465,0.0010898539,0.00021896689,0.0007421959,0.0007652487,0.00033709596],"domain_scores_gemma":[0.99160916,0.0064412286,0.0006417105,0.0004390461,0.00068605074,0.00018272185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003514334,0.0017484174,0.0019308976,0.0017515086,0.0010154002,0.0017977769,0.002070202,0.0028607869,0.0035522762],"category_scores_gemma":[0.011278451,0.0010269505,0.0016111004,0.0020164184,0.0013647187,0.0029190942,0.0017102862,0.0023173452,0.00070821407],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044642293,0.000517481,0.004016393,0.0007516705,0.00024135248,0.0005656575,0.00042689923,0.54844266,0.014250077,0.03315549,0.008935384,0.38825047],"study_design_scores_gemma":[0.00012003621,0.0002453311,0.001303557,0.000073005685,0.00012717526,0.00038234552,0.00016524257,0.93787396,0.0063566654,0.04796789,0.005343416,0.00004148487],"about_ca_topic_score_codex":0.0030359293,"about_ca_topic_score_gemma":0.0019995712,"teacher_disagreement_score":0.0035522762,"about_ca_system_score_codex":0.0014597393,"about_ca_system_score_gemma":0.0016087367,"threshold_uncertainty_score":0.018585801},"labels":[],"label_agreement":null},{"id":"W2017243662","doi":"10.1155/2010/932686","title":"A Tester-Assisted Methodology for Test Redundancy Detection","year":2010,"lang":"en","type":"article","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Test suite; Computer science; Redundancy (engineering); Reliability engineering; Fault detection and isolation; Suite; Test (biology); Test case; Code coverage; Data mining; Java; Software; Machine learning; Artificial intelligence; Programming language; Operating system; Engineering","score_opus":0.023343681542986727,"score_gpt":0.310099371844932,"score_spread":0.2867556903019453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2017243662","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007196711,0.00005989945,0.98886347,0.000043627344,0.000017209484,0.00017594568,0.00005952293,0.0030756327,0.0005079082],"genre_scores_gemma":[0.12284039,0.00006613692,0.8743673,0.000088218345,0.000021689088,0.00043288356,0.00044674936,0.00036695265,0.001369721],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99636465,0.0013764818,0.0002266313,0.00047920592,0.0014299682,0.00012311676],"domain_scores_gemma":[0.99207413,0.0034353973,0.000853789,0.0016270431,0.0018882769,0.00012127413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022019416,0.0012035829,0.0008961337,0.0027537972,0.00036651842,0.00084969,0.0019591695,0.0010428672,0.0020306865],"category_scores_gemma":[0.011341791,0.00047610546,0.0011166018,0.0011462611,0.00057067786,0.00072682963,0.0007936649,0.001298115,0.0009773527],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025401235,0.00045418247,0.005341349,0.00062632037,0.0002071457,0.0009281395,0.00048742213,0.044235207,0.19268139,0.015445775,0.004069336,0.73526967],"study_design_scores_gemma":[0.00012359742,0.0009080487,0.004380014,0.00008989048,0.00018283175,0.003749826,0.000081446145,0.8111007,0.14993544,0.013096356,0.016231116,0.00012071194],"about_ca_topic_score_codex":0.00086457736,"about_ca_topic_score_gemma":0.0009450206,"teacher_disagreement_score":0.0027537972,"about_ca_system_score_codex":0.0003963025,"about_ca_system_score_gemma":0.0013488998,"threshold_uncertainty_score":0.011645138},"labels":[],"label_agreement":null},{"id":"W2017616266","doi":"10.1109/wcre.2011.20","title":"Reverse Engineering Co-maintenance Relationships Using Conceptual Analysis of Source Code","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Cohesion (chemistry); Source code; Software maintenance; Visualization; Reverse engineering; Conceptual model; Software engineering; Data science; Programming language; Data mining; Database; Software development","score_opus":0.1026566769470182,"score_gpt":0.2796243366358938,"score_spread":0.17696765968887562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2017616266","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27389786,0.0010975411,0.71750355,0.0006105739,0.000053160285,0.00022327763,0.00059866725,0.0022098206,0.0038055875],"genre_scores_gemma":[0.7255107,0.00038950078,0.27156478,0.000048015692,0.00003277491,0.00016011308,0.0010066729,0.00038591205,0.00090155564],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9954224,0.0015162447,0.00028117993,0.0009770272,0.001646128,0.00015697436],"domain_scores_gemma":[0.92120683,0.052580085,0.008940408,0.009325848,0.0073902896,0.0005565928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055962475,0.0006682567,0.0004831181,0.007815659,0.0009796199,0.0034652739,0.0013101235,0.00084514875,0.001694821],"category_scores_gemma":[0.05715683,0.0005628357,0.00075127016,0.0058906362,0.0011408775,0.0072851274,0.0021124692,0.001568377,0.00035289818],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039885426,0.00031872786,0.17935145,0.0011325831,0.00041237328,0.0011300805,0.02180796,0.030999182,0.02639665,0.07121167,0.003337932,0.6635025],"study_design_scores_gemma":[0.000074623385,0.0004618611,0.13573818,0.00050689926,0.0004512348,0.0019332379,0.009384767,0.65434605,0.032754727,0.12801617,0.03605229,0.00027993662],"about_ca_topic_score_codex":0.004516288,"about_ca_topic_score_gemma":0.007097638,"teacher_disagreement_score":0.007815659,"about_ca_system_score_codex":0.00128553,"about_ca_system_score_gemma":0.0012818316,"threshold_uncertainty_score":0.02959615},"labels":[],"label_agreement":null},{"id":"W2017671454","doi":"10.1109/fosm.2008.4659256","title":"The past, present, and future of software evolution","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":128,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Waterloo","funders":"","keywords":"Software evolution; Software development; Computer science; Software analytics; Social software engineering; Software; Software engineering; Software system; Software maintenance; Software construction; Data science; Programming language","score_opus":0.012369626224032822,"score_gpt":0.23436887583009205,"score_spread":0.22199924960605924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2017671454","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12831083,0.39069068,0.119550906,0.22612137,0.003239384,0.00006386233,0.00032798137,0.00024513312,0.1314499],"genre_scores_gemma":[0.8374245,0.12554267,0.023381658,0.004139588,0.002932773,0.00007594479,0.0001433288,0.00005883358,0.0063006547],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972982,0.0012447627,0.00014838258,0.00042806234,0.0007072239,0.00017342932],"domain_scores_gemma":[0.99560106,0.0024075557,0.0004747133,0.00043479423,0.00076559297,0.00031618561],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003962863,0.0003238975,0.00042620502,0.0013838039,0.0017619175,0.005929433,0.00068944495,0.0029463784,0.0019508287],"category_scores_gemma":[0.0082033435,0.00033397597,0.0003189882,0.0019130827,0.0076722465,0.017194355,0.0017389562,0.0035070032,0.00031989822],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046752284,0.00003249471,0.0040604006,0.00033451553,0.000016172049,0.00012133133,0.0021617967,0.001340467,0.00079210295,0.8392043,0.003125384,0.14876431],"study_design_scores_gemma":[0.000010940489,0.000063650514,0.005298423,0.0006915543,0.000023907396,0.00053152104,0.0023986062,0.006210257,0.0005024078,0.83093715,0.1532707,0.000060790095],"about_ca_topic_score_codex":0.0016198038,"about_ca_topic_score_gemma":0.0013040868,"teacher_disagreement_score":0.005929433,"about_ca_system_score_codex":0.0022329097,"about_ca_system_score_gemma":0.0017141796,"threshold_uncertainty_score":0.020957828},"labels":[],"label_agreement":null},{"id":"W2017993165","doi":"10.1016/j.advengsoft.2009.07.003","title":"Horizontal dispersion of software functional size with IFPUG and COSMIC units","year":2009,"lang":"en","type":"article","venue":"Advances in Engineering Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; École de Technologie Supérieure","funders":"","keywords":"Function point; Software; Standardization; Computer science; Variable (mathematics); Dispersion (optics); Software development; Function (biology); Observational error; Software metric; Reliability engineering; Software quality; Statistics; Mathematics; Engineering; Physics; Programming language; Optics","score_opus":0.005946882834828772,"score_gpt":0.21380460932094097,"score_spread":0.2078577264861122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2017993165","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95542026,0.0005109951,0.027860934,0.00023136573,0.00005427899,0.000018515411,0.0014596028,0.0011318417,0.013312123],"genre_scores_gemma":[0.9963887,0.000031772714,0.0023650273,0.000017455162,0.000017443313,0.000011225363,0.00058728224,0.00010776837,0.00047336394],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99832135,0.00044420274,0.000085555665,0.00042355128,0.00050209445,0.00022319435],"domain_scores_gemma":[0.97108185,0.017278235,0.0035516608,0.004775772,0.0026796767,0.0006327826],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026842055,0.0003620949,0.0003960746,0.003922585,0.00037150833,0.0013837479,0.00078227837,0.0008638313,0.003289271],"category_scores_gemma":[0.030042239,0.00027860518,0.00044804259,0.0038681633,0.00090019236,0.0017671636,0.0009857946,0.00088210066,0.00054885494],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008662194,0.00013371401,0.5907428,0.0001739571,0.00037572152,0.00041146184,0.0018639065,0.18187837,0.0063429684,0.057378948,0.0081331115,0.15169883],"study_design_scores_gemma":[0.000039722632,0.00025196432,0.5250539,0.0001392646,0.00014314604,0.0007186843,0.0009867314,0.4159484,0.008102815,0.041678935,0.006830643,0.00010574301],"about_ca_topic_score_codex":0.0055654,"about_ca_topic_score_gemma":0.0039758673,"teacher_disagreement_score":0.0055654,"about_ca_system_score_codex":0.0009700929,"about_ca_system_score_gemma":0.00035868105,"threshold_uncertainty_score":0.014195621},"labels":[],"label_agreement":null},{"id":"W2018392829","doi":"10.1145/2579281.2579311","title":"Technical debt at the crossroads of research and practice","year":2014,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Metaphor; Software; Debt; Key (lock); Quality (philosophy); Computer science; Engineering management; Empirical research; Software quality; Software engineering; Engineering; Software development; Knowledge management; Business; Computer security; Finance","score_opus":0.03772669070019714,"score_gpt":0.3336635637803617,"score_spread":0.2959368730801646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2018392829","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05923222,0.10439424,0.085921004,0.58787537,0.0042551504,0.00023108884,0.00016223539,0.00040088999,0.15752785],"genre_scores_gemma":[0.9057005,0.026763465,0.026500614,0.02878224,0.0028816042,0.0005707759,0.000103843515,0.0004801565,0.008216759],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8978823,0.06744523,0.0051688696,0.008186119,0.017110744,0.004206724],"domain_scores_gemma":[0.77002716,0.17756082,0.007325519,0.019829929,0.015749607,0.009506887],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.12362633,0.0008911323,0.0020113813,0.012485653,0.010763558,0.04744744,0.0033409707,0.011178827,0.008015739],"category_scores_gemma":[0.16265719,0.0017375321,0.00072410255,0.010928346,0.085124366,0.084491834,0.0287971,0.02030608,0.0014149575],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000050819803,0.0000713135,0.0017559371,0.00062290055,0.000047304937,0.00020848363,0.042196847,0.00041392032,0.00031276033,0.89123875,0.007261581,0.055819504],"study_design_scores_gemma":[0.00003930289,0.00007483571,0.001510241,0.0024687157,0.000020797146,0.00019377004,0.025341175,0.0005881942,0.00021815552,0.8721692,0.0973196,0.000055962402],"about_ca_topic_score_codex":0.004162422,"about_ca_topic_score_gemma":0.0024068833,"teacher_disagreement_score":0.87637365,"about_ca_system_score_codex":0.017042011,"about_ca_system_score_gemma":0.017533502,"threshold_uncertainty_score":0.6538063},"labels":[],"label_agreement":null},{"id":"W2018414565","doi":"10.1145/2491509.2491515","title":"The value of design rationale information","year":2013,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Norges Forskningsråd","keywords":"Documentation; Computer science; Personalization; Context (archaeology); Value (mathematics); World Wide Web; Programming language","score_opus":0.07125196779048361,"score_gpt":0.29658960558982694,"score_spread":0.22533763779934335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2018414565","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9223481,0.0015379203,0.03804011,0.002433552,0.00006721656,0.00068337447,0.00020419086,0.00044045644,0.034245078],"genre_scores_gemma":[0.9682442,0.00028367015,0.02976554,0.00024249089,0.000027795524,0.00015585298,0.00016439674,0.000053847434,0.0010622034],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9235763,0.04606177,0.00542973,0.0028649922,0.020977318,0.0010897886],"domain_scores_gemma":[0.4092968,0.4888514,0.032584462,0.051220287,0.0149747655,0.0030723219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.037581634,0.00077129214,0.0005875738,0.0026445729,0.00075759285,0.0065659573,0.0018731083,0.001932398,0.0018744299],"category_scores_gemma":[0.27770883,0.0008141117,0.00062476646,0.00185092,0.0018044723,0.0077359886,0.0027571688,0.0024581475,0.0004810723],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017737579,0.0019317608,0.09612615,0.002201723,0.0003846243,0.0005765003,0.02058841,0.0071123824,0.034550276,0.017607007,0.0019906084,0.8151569],"study_design_scores_gemma":[0.0013483647,0.017070435,0.5074738,0.006483981,0.0026105302,0.0060781417,0.02708307,0.08021771,0.0811794,0.13137959,0.13786675,0.0012081778],"about_ca_topic_score_codex":0.00058074313,"about_ca_topic_score_gemma":0.00078059395,"teacher_disagreement_score":0.037581634,"about_ca_system_score_codex":0.0022782334,"about_ca_system_score_gemma":0.0024384428,"threshold_uncertainty_score":0.19875306},"labels":[],"label_agreement":null},{"id":"W2018718931","doi":"10.1109/vissoft.2014.26","title":"Templated Visualization of Object State with Vebugger","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Object (grammar); Java; Representation (politics); Debugging; Visualization; Set (abstract data type); Programming language; State (computer science); Hierarchy; Listing (finance); Template; Information retrieval; Artificial intelligence; Computer graphics (images)","score_opus":0.010135333212106015,"score_gpt":0.2608787933400095,"score_spread":0.2507434601279035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2018718931","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016617345,0.00027810707,0.7878341,0.000308582,0.00021530758,0.0001719526,0.0042322655,0.18137158,0.008970748],"genre_scores_gemma":[0.195227,0.00062121666,0.7374209,0.0005231523,0.00007142339,0.00069957227,0.010921096,0.04007579,0.014439936],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994221,0.00011490447,0.00005569011,0.00010996828,0.00023549159,0.00006191446],"domain_scores_gemma":[0.9968514,0.0016185824,0.00014588653,0.0007456368,0.0004538537,0.00018470702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014694962,0.0012080974,0.0007420555,0.0018659934,0.0004074387,0.0023031346,0.0015953659,0.0011951751,0.018165456],"category_scores_gemma":[0.005464043,0.00068659335,0.00087328407,0.0010712302,0.0003986166,0.0024620302,0.0022371185,0.0018058581,0.0043588523],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025447102,0.00068398897,0.009957757,0.0017424143,0.00023089437,0.0014835027,0.0047987853,0.03314545,0.09169133,0.042928763,0.24286857,0.56792384],"study_design_scores_gemma":[0.00047203177,0.0003374907,0.008473885,0.0005888668,0.00011814572,0.0016870891,0.0004255577,0.34529823,0.17116404,0.03337696,0.4376052,0.0004525698],"about_ca_topic_score_codex":0.0027433762,"about_ca_topic_score_gemma":0.002752678,"teacher_disagreement_score":0.018165456,"about_ca_system_score_codex":0.0005071464,"about_ca_system_score_gemma":0.0006803488,"threshold_uncertainty_score":0.060769558},"labels":[],"label_agreement":null},{"id":"W2018747777","doi":"10.1109/vissoft.2014.19","title":"Validation of Software Visualization Tools: A Systematic Mapping Study","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Visualization; Software visualization; Computer science; Software; Data science; Software engineering; Verification and validation; Software development; Information visualization; Software construction; Data mining; Engineering","score_opus":0.0382877747461992,"score_gpt":0.29489393087187676,"score_spread":0.25660615612567755,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2018747777","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8047158,0.0906874,0.06372499,0.0022155424,0.0002697346,0.025351109,0.003003181,0.00025400845,0.00977821],"genre_scores_gemma":[0.8951843,0.01963717,0.06444635,0.0006084611,0.00006616329,0.01778047,0.001457697,0.0001234778,0.0006959355],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.77683437,0.13065004,0.04422021,0.009356492,0.03615524,0.002783565],"domain_scores_gemma":[0.32027408,0.54222184,0.048134904,0.021505304,0.066458575,0.001405308],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.18103038,0.0015319969,0.0030757748,0.0553082,0.0037539208,0.005301294,0.0021830923,0.002034815,0.0013531103],"category_scores_gemma":[0.43782154,0.0015830095,0.0035045142,0.033798303,0.004743687,0.008782594,0.00687519,0.0015737024,0.00034296015],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070968264,0.0006820559,0.13656989,0.13055089,0.004391855,0.002122169,0.27033147,0.0014840367,0.004376087,0.0072031054,0.002917643,0.43866107],"study_design_scores_gemma":[0.0006943506,0.0041867434,0.19371597,0.29180047,0.011999609,0.0049237227,0.36860004,0.0061442386,0.017117482,0.01186879,0.088310726,0.0006378647],"about_ca_topic_score_codex":0.0037134152,"about_ca_topic_score_gemma":0.005944883,"teacher_disagreement_score":0.18103038,"about_ca_system_score_codex":0.0074343975,"about_ca_system_score_gemma":0.022921339,"threshold_uncertainty_score":0.95739156},"labels":[],"label_agreement":null},{"id":"W2018873511","doi":"10.1145/1028174.971429","title":"Developing principles of GUI programming using views","year":2004,"lang":"en","type":"article","venue":"ACM SIGCSE Bulletin","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Microsoft Research","keywords":"Computer science; Programming language; Programmer; XML; Notation; Java; Event-driven programming; Programming paradigm; Software engineering; Procedural programming; Inductive programming; Operating system","score_opus":0.1032978941526476,"score_gpt":0.32310445525516995,"score_spread":0.21980656110252234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2018873511","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00096317375,0.00015083664,0.9930884,0.00042917146,0.000047511676,0.000066102344,0.000019939196,0.001205655,0.0040292614],"genre_scores_gemma":[0.042089622,0.0006037574,0.9508739,0.0004378209,0.000119561584,0.0002915921,0.000131965,0.0008378836,0.0046138726],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9955379,0.0017555986,0.0004079626,0.0006759095,0.0013079727,0.00031469308],"domain_scores_gemma":[0.9929311,0.0034037495,0.0003596276,0.0016269241,0.0012792259,0.00039935438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00888173,0.00094949524,0.0006757984,0.0011590873,0.0012916869,0.0053869947,0.0035002006,0.0017230223,0.003557259],"category_scores_gemma":[0.010713975,0.0016689061,0.0018565588,0.0006505173,0.005669529,0.007994583,0.0049053845,0.0064699943,0.001922696],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028012182,0.000037373793,0.00054938765,0.00019213487,0.00003119256,0.00017052119,0.0019988003,0.0046825404,0.0020992444,0.9207198,0.005025815,0.06446501],"study_design_scores_gemma":[0.00007061402,0.0000939999,0.0002603702,0.00028899015,0.000047461886,0.0005141867,0.0004166734,0.03979129,0.0065000625,0.7594175,0.19252288,0.00007606647],"about_ca_topic_score_codex":0.0020762032,"about_ca_topic_score_gemma":0.002131926,"teacher_disagreement_score":0.00888173,"about_ca_system_score_codex":0.0011701777,"about_ca_system_score_gemma":0.002269554,"threshold_uncertainty_score":0.04697168},"labels":[],"label_agreement":null},{"id":"W2018890516","doi":"10.1109/icsm.2011.6080794","title":"Late propagation in software clones","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Synchronizing; Software evolution; Biology; Software maintenance; Computer science; Software system; Software; Genetics; Programming language; Gene","score_opus":0.03403372786823052,"score_gpt":0.2487466589834545,"score_spread":0.21471293111522397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2018890516","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6542934,0.0017738284,0.33769014,0.00035505506,0.00006811449,0.0002457336,0.000109988825,0.001292923,0.004170833],"genre_scores_gemma":[0.94998866,0.0003396389,0.046278324,0.00017958306,0.00005012747,0.00010177843,0.00014586387,0.00018049082,0.0027354541],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9926714,0.0015966623,0.00066773425,0.0012792993,0.0031854226,0.0005994314],"domain_scores_gemma":[0.8912003,0.058283854,0.020426007,0.015947282,0.012746735,0.0013958688],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004070124,0.0006946129,0.0007019028,0.0027586783,0.0011212656,0.0020757804,0.0014157959,0.0018056115,0.000874706],"category_scores_gemma":[0.0525745,0.0006992369,0.0008741799,0.0017451572,0.0021428608,0.004628657,0.0022123018,0.0017178219,0.00022918219],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012122273,0.0004954016,0.31498745,0.0010768948,0.00038795397,0.0056746183,0.008547806,0.057540912,0.08969303,0.089667976,0.0026291858,0.42808667],"study_design_scores_gemma":[0.00027002257,0.0036495137,0.15395668,0.0006456536,0.0010901397,0.031043397,0.0024181253,0.41801807,0.16076986,0.18723227,0.040460095,0.00044611082],"about_ca_topic_score_codex":0.0015795572,"about_ca_topic_score_gemma":0.0011870486,"teacher_disagreement_score":0.004070124,"about_ca_system_score_codex":0.0010833322,"about_ca_system_score_gemma":0.0009872072,"threshold_uncertainty_score":0.021525085},"labels":[],"label_agreement":null},{"id":"W2019003365","doi":"10.1109/isie.2006.296135","title":"An ISO/IEC standards-based quality requirement definition approach: Applicative analysis of three quality requirements definition methods","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Quality (philosophy); Requirements engineering; Quality of analytical results; Requirements analysis; Non-functional requirement; Requirements management; Software quality; Requirement; Identification (biology); Software quality control; Software requirements; Requirement prioritization; Business requirements; Software engineering; Software requirements specification; Non-functional testing; Task (project management); Systems engineering; Quality assurance; Software; Quality policy; Software development; Engineering; Business process; Software construction; Work in process; Operations management","score_opus":0.20856793086436992,"score_gpt":0.44734346673647013,"score_spread":0.2387755358721002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2019003365","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0068303584,0.0004965174,0.9801769,0.00076395227,0.00006268597,0.00094834843,0.00011626008,0.00034276207,0.010262173],"genre_scores_gemma":[0.04760769,0.00041594822,0.94915056,0.00017307259,0.000024784731,0.0012168599,0.00037818323,0.00008980095,0.0009431792],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.95505583,0.015629003,0.0044229203,0.0015528769,0.02254813,0.0007912047],"domain_scores_gemma":[0.94888765,0.022221172,0.0043429853,0.003477637,0.020589987,0.0004805846],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03132132,0.0013757147,0.0008803324,0.009401792,0.0012506078,0.0049669067,0.002584958,0.002103486,0.001655408],"category_scores_gemma":[0.04995685,0.00059553137,0.0026178479,0.00614262,0.0024159062,0.0047301757,0.002219587,0.0022831724,0.0006605057],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002383162,0.000820526,0.010415857,0.0023878906,0.00027341617,0.00040938845,0.005458172,0.026772108,0.015694426,0.41408268,0.007661294,0.51578593],"study_design_scores_gemma":[0.0003247058,0.0017845822,0.026955819,0.0045918445,0.00073400594,0.0024226783,0.0075398595,0.3834918,0.04026946,0.33298063,0.1982728,0.0006318222],"about_ca_topic_score_codex":0.0030400788,"about_ca_topic_score_gemma":0.002605183,"teacher_disagreement_score":0.03132132,"about_ca_system_score_codex":0.0045467764,"about_ca_system_score_gemma":0.008471634,"threshold_uncertainty_score":0.16564494},"labels":[],"label_agreement":null},{"id":"W2019009588","doi":"10.4304/jsw.8.2.327-336","title":"Qualitative Analysis for the Impact of Accounting for Special Methods in Object-Oriented Class Cohesion Measurement","year":2013,"lang":"en","type":"article","venue":"Journal of Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Kuwait University; University of Alberta","keywords":"Computer science; Cohesion (chemistry); Class (philosophy); Artificial intelligence","score_opus":0.0869237389609992,"score_gpt":0.437746582337826,"score_spread":0.3508228433768268,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2019009588","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92776066,0.00029858213,0.06522128,0.0005091685,0.000027169222,0.00013137922,0.00012956966,0.00012234917,0.005799784],"genre_scores_gemma":[0.9930482,0.000029516164,0.0067211837,0.000015062662,0.000002193795,0.00003460567,0.000023149501,0.0000080831405,0.00011787249],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9857503,0.00778222,0.00064621243,0.00052540837,0.0049316515,0.00036422166],"domain_scores_gemma":[0.80528104,0.15758881,0.01216026,0.00545264,0.018664867,0.0008523744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013627366,0.00023076113,0.00023646228,0.0020676425,0.0006193158,0.0012888819,0.000549553,0.00041969772,0.0014124742],"category_scores_gemma":[0.0692739,0.00016333087,0.0004108417,0.0016728705,0.0016697334,0.0015433664,0.00094124826,0.0005882843,0.00010383729],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016207494,0.0007376403,0.41349375,0.0019544777,0.00038293644,0.0009796295,0.024737582,0.053220425,0.097795784,0.05653303,0.0030313937,0.34551266],"study_design_scores_gemma":[0.00009415806,0.0025745176,0.49767908,0.0008129696,0.00055705494,0.00072194013,0.028672244,0.28837928,0.13250299,0.03613864,0.0116167,0.000250383],"about_ca_topic_score_codex":0.0023593835,"about_ca_topic_score_gemma":0.002905881,"teacher_disagreement_score":0.013627366,"about_ca_system_score_codex":0.0024509882,"about_ca_system_score_gemma":0.0013641638,"threshold_uncertainty_score":0.07206923},"labels":[],"label_agreement":null},{"id":"W2019257047","doi":"10.1007/s10664-015-9381-9","title":"An empirical study of the impact of modern code review practices on software quality","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":320,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Code review; Software quality; Computer science; Software engineering; Software quality analyst; Software inspection; Software peer review; Software quality management; Static program analysis; Software construction; Software quality assurance; Software development; Software; Programming language","score_opus":0.14670699522306,"score_gpt":0.4503907650299953,"score_spread":0.30368376980693534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2019257047","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9981111,0.00033354427,0.00028904053,0.00018904688,0.0000063972707,0.000025702784,0.00003518358,0.00001271137,0.0009971929],"genre_scores_gemma":[0.99908066,0.00013081003,0.00047691743,0.00004606325,0.000014801438,0.000014533256,0.00003727883,0.0000051808456,0.00019381648],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9840423,0.007693855,0.0011451403,0.0011252231,0.005061395,0.00093203847],"domain_scores_gemma":[0.42364755,0.4192705,0.10102864,0.013889635,0.03242593,0.009737784],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014670518,0.0002526799,0.0002813511,0.0031193113,0.0009291299,0.0018738813,0.0010626176,0.0009039212,0.0019077983],"category_scores_gemma":[0.171371,0.00037137713,0.0004334131,0.0032487742,0.0017555343,0.002449424,0.001432291,0.0018614616,0.00020268641],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013689232,0.006450636,0.8537005,0.00059933175,0.0003711723,0.0002747418,0.005974135,0.0022331676,0.0037640866,0.0018951807,0.0011070912,0.122261055],"study_design_scores_gemma":[0.0001277591,0.0031404078,0.9865919,0.00015017923,0.00017759226,0.00021836744,0.0034837709,0.0027701317,0.0013104846,0.0005239119,0.0014712424,0.000034276116],"about_ca_topic_score_codex":0.0075919926,"about_ca_topic_score_gemma":0.013931158,"teacher_disagreement_score":0.9853295,"about_ca_system_score_codex":0.003458273,"about_ca_system_score_gemma":0.0042577637,"threshold_uncertainty_score":0.077586055},"labels":[],"label_agreement":null},{"id":"W2019441277","doi":"10.1109/icsm.2011.6080802","title":"Evaluating software clustering using multiple simulated authoritative decompositions","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Computer science; Data mining; Software; Decomposition; Correlation clustering; CURE data clustering algorithm; Machine learning; Artificial intelligence; Programming language","score_opus":0.1972400476707733,"score_gpt":0.39314260852102956,"score_spread":0.19590256085025626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2019441277","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8615144,0.00016917053,0.1339079,0.000111461326,0.000042860232,0.00025418377,0.00023890703,0.00081960147,0.0029415889],"genre_scores_gemma":[0.90881276,0.00005965017,0.08982842,0.000025069377,0.000013578284,0.00023937858,0.0005458492,0.00007770662,0.00039756446],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9867091,0.008229773,0.00074663793,0.0009829288,0.0029308207,0.000400623],"domain_scores_gemma":[0.93528825,0.040842697,0.0046003163,0.007993834,0.009999825,0.001274975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014146913,0.0014261303,0.0012333379,0.0040719626,0.0010656158,0.0019872985,0.0015963309,0.0014235523,0.0008385086],"category_scores_gemma":[0.046306767,0.00056725135,0.0009785369,0.002692859,0.0011442384,0.0021978752,0.0017788649,0.00091193273,0.0002478139],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012046545,0.00085360894,0.023961715,0.0002165525,0.00039142685,0.00013455529,0.0007310128,0.8892707,0.006664656,0.006126646,0.0011694621,0.069275014],"study_design_scores_gemma":[0.00006179497,0.00051053445,0.0024518946,0.000017712611,0.000039400795,0.00006269206,0.00018254429,0.9862434,0.0077531403,0.002216985,0.00043085878,0.000029104984],"about_ca_topic_score_codex":0.0021212334,"about_ca_topic_score_gemma":0.0032641988,"teacher_disagreement_score":0.014146913,"about_ca_system_score_codex":0.0024389324,"about_ca_system_score_gemma":0.0010927931,"threshold_uncertainty_score":0.07481688},"labels":[],"label_agreement":null},{"id":"W2019628947","doi":"10.1109/icsm.2012.6405285","title":"Models are code too: Near-miss clone detection for Simulink models","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":74,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; General Motors of Canada; Technische Universität München","keywords":"Computer science; Source code; clone (Java method); Code (set theory); Matching (statistics); Graph; Detector; Programming language; Graphical model; Identification (biology); Theoretical computer science; Artificial intelligence; Mathematics","score_opus":0.06396359254443656,"score_gpt":0.29260169565127925,"score_spread":0.22863810310684268,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2019628947","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.121467836,0.00009186671,0.8683595,0.0001246563,0.000018515891,0.00007012004,0.0002656878,0.008441383,0.0011604588],"genre_scores_gemma":[0.663051,0.0000970809,0.33306825,0.00008155308,0.00000733774,0.000080873826,0.0007928471,0.00089890667,0.0019220267],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985145,0.00028762873,0.000077025536,0.00032539715,0.000728836,0.0000666392],"domain_scores_gemma":[0.99370086,0.0026607688,0.0011700301,0.0015970635,0.0007259999,0.00014530387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00080842664,0.00079815777,0.00056212157,0.0011914456,0.00042794237,0.0007576286,0.0010527772,0.0009646156,0.0013252995],"category_scores_gemma":[0.013307924,0.0004874097,0.00069771515,0.000773298,0.0007403797,0.0020308874,0.0011968001,0.0009117498,0.000418565],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010347946,0.00024577964,0.034113508,0.0005345857,0.00020946098,0.0024735264,0.0017809024,0.49350727,0.09227293,0.048128374,0.004158685,0.32154015],"study_design_scores_gemma":[0.000018912542,0.00009653443,0.0014258451,0.000020847907,0.000035044908,0.00035120305,0.00010311771,0.9425472,0.034728345,0.01726609,0.0033837366,0.00002305992],"about_ca_topic_score_codex":0.00333869,"about_ca_topic_score_gemma":0.005526478,"teacher_disagreement_score":0.00333869,"about_ca_system_score_codex":0.0006310341,"about_ca_system_score_gemma":0.00082442083,"threshold_uncertainty_score":0.006638527},"labels":[],"label_agreement":null},{"id":"W2020166315","doi":"10.1145/1137983.1138020","title":"Using evolutionary annotations from change logs to enhance program comprehension","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Source code; Computer science; Program comprehension; Workbench; Software evolution; Eclipse; Programming language; Software maintenance; Code (set theory); Software; Filter (signal processing); Evolutionary algorithm; Software engineering; Artificial intelligence; Software development; Software system; Set (abstract data type); Software construction; Visualization","score_opus":0.06159014107737977,"score_gpt":0.351656608422654,"score_spread":0.29006646734527425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2020166315","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11840962,0.00025668708,0.79272187,0.0010077289,0.00012417856,0.00071416365,0.0024050076,0.07979839,0.004562288],"genre_scores_gemma":[0.27162442,0.00024237185,0.7128579,0.00019114978,0.0001067073,0.0004858943,0.006310212,0.0045851935,0.0035960986],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9952632,0.0018592622,0.00039789415,0.000891829,0.0014210523,0.000166769],"domain_scores_gemma":[0.8755213,0.09018888,0.008444573,0.012723096,0.011972092,0.0011499313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007976924,0.0018818016,0.001141607,0.0057219174,0.00087345426,0.0033483207,0.002084737,0.0020109967,0.004024949],"category_scores_gemma":[0.0777107,0.0010473636,0.0006503239,0.0028021713,0.0007392981,0.007963019,0.0022741908,0.002345202,0.0017529399],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011622426,0.001166128,0.026003728,0.0013409301,0.00011366041,0.0016232668,0.014717058,0.012786077,0.04550614,0.005883144,0.015035668,0.87466204],"study_design_scores_gemma":[0.00047861313,0.0010930892,0.03716521,0.00073701947,0.00036127755,0.0021649578,0.0037111603,0.6595663,0.16451947,0.034129653,0.09545571,0.0006175719],"about_ca_topic_score_codex":0.0027791415,"about_ca_topic_score_gemma":0.004157132,"teacher_disagreement_score":0.007976924,"about_ca_system_score_codex":0.0007245102,"about_ca_system_score_gemma":0.0015535591,"threshold_uncertainty_score":0.0421865},"labels":[],"label_agreement":null},{"id":"W2020259809","doi":"10.1109/qsic.2012.17","title":"Strategic Management of Technical Debt: Tutorial Synopsis","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Carnegie Mellon University; University of Pittsburgh","keywords":"Technical debt; Debt; Computer science; Internal debt; Leverage (statistics); Business; Risk analysis (engineering); Finance; Software development; Software","score_opus":0.03076217188995573,"score_gpt":0.28600483120758863,"score_spread":0.2552426593176329,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2020259809","genre_codex":"review","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042816577,0.6836123,0.08585495,0.018522637,0.010781295,0.0004160113,0.0004488483,0.00044883226,0.19563349],"genre_scores_gemma":[0.028752934,0.8236619,0.037403226,0.008271781,0.0133170765,0.0006049926,0.0006259282,0.00032465547,0.0870374],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995179,0.00014851936,0.000058670255,0.00007051881,0.00014737903,0.00005705278],"domain_scores_gemma":[0.99910223,0.0005089769,0.0000791909,0.000031174543,0.00020361824,0.00007476209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009027597,0.0012903743,0.0007296521,0.0034193224,0.0010614758,0.0039902497,0.0009807419,0.0024384558,0.0123955645],"category_scores_gemma":[0.0020367585,0.0006097398,0.0006638634,0.0036294798,0.001403117,0.0070973425,0.0017423715,0.0032270518,0.0053758463],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000056111247,0.0002660525,0.00082816294,0.0042632413,0.000031440006,0.0008212097,0.0024710423,0.0040082857,0.002303747,0.32118234,0.27992216,0.3838462],"study_design_scores_gemma":[0.0000039434303,0.00006180334,0.00082153274,0.0017120761,0.000008124758,0.00068368483,0.0005086731,0.0009095833,0.00028767835,0.035333645,0.9596384,0.000030885683],"about_ca_topic_score_codex":0.0030095093,"about_ca_topic_score_gemma":0.003649075,"teacher_disagreement_score":0.0123955645,"about_ca_system_score_codex":0.0024269107,"about_ca_system_score_gemma":0.0013327409,"threshold_uncertainty_score":0.04146731},"labels":[],"label_agreement":null},{"id":"W2020316880","doi":"10.1145/1039174.1039188","title":"Report on MSR 2004","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Presentation (obstetrics); Session (web analytics); Computer science; Software; Reuse; Process (computing); Software engineering; Software development; World Wide Web; Data science; Engineering","score_opus":0.01739211010709344,"score_gpt":0.26033539130103944,"score_spread":0.242943281193946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2020316880","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046606823,0.008024884,0.003703496,0.025719611,0.050539244,0.0014278915,0.05264387,0.006454008,0.8468265],"genre_scores_gemma":[0.006029316,0.002258188,0.0013788218,0.0028011943,0.0036885026,0.00029971657,0.02359332,0.0011400101,0.95881104],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961637,0.0004911877,0.00029391283,0.00058214134,0.0020268674,0.0004422201],"domain_scores_gemma":[0.992607,0.0004746372,0.00026717113,0.0010452305,0.00418386,0.0014221697],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0041262526,0.0012564713,0.0011828459,0.0035397022,0.0018019602,0.0069251372,0.0024171567,0.0026340343,0.5287401],"category_scores_gemma":[0.008961292,0.0004186799,0.0010611892,0.0025800057,0.00038396654,0.0029641027,0.003159067,0.0020210384,0.54978573],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017675271,0.00006548233,0.00028896568,0.00014153332,0.000007829323,0.000074681324,0.00004935654,0.00005652439,0.00048001128,0.0017729787,0.9466002,0.05028565],"study_design_scores_gemma":[0.000012047368,0.000036872683,0.00042852268,0.000041702962,0.0000029972346,0.000024929705,0.000033215,0.000029433511,0.00021615835,0.0002327318,0.99893695,0.0000043252808],"about_ca_topic_score_codex":0.0040291487,"about_ca_topic_score_gemma":0.0053049964,"teacher_disagreement_score":0.5287401,"about_ca_system_score_codex":0.003053225,"about_ca_system_score_gemma":0.0033868705,"threshold_uncertainty_score":0.672195},"labels":[],"label_agreement":null},{"id":"W2020620746","doi":"10.1007/s00500-007-0215-6","title":"On the possibilities of (pseudo-) software cloning from external interactions","year":2007,"lang":"en","type":"article","venue":"Soft Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Cloning (programming); Computer science; Software; Software engineering; Programming language","score_opus":0.02131087380848516,"score_gpt":0.2869729112439359,"score_spread":0.26566203743545075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2020620746","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059417017,0.0008941346,0.8382266,0.0034595018,0.0004065297,0.00009042383,0.00008433084,0.0007255758,0.09669595],"genre_scores_gemma":[0.79584694,0.0013080197,0.17631108,0.0010200186,0.0002823169,0.0002588732,0.0001578594,0.0006992597,0.024115644],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9934155,0.003021852,0.00027154307,0.0007720945,0.0018606537,0.0006583672],"domain_scores_gemma":[0.9727912,0.014942875,0.0008290406,0.009606731,0.0013243975,0.00050575176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005571267,0.00087120384,0.00093167054,0.0011768439,0.003266741,0.0068833837,0.0027513823,0.0036924004,0.0073026847],"category_scores_gemma":[0.030687071,0.0010071918,0.0018979865,0.0011442864,0.011796978,0.014414255,0.0076399753,0.0052600564,0.0018124154],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020212417,0.00001062942,0.00014541496,0.000020838203,0.0000043770365,0.00007072881,0.00023256909,0.0013442351,0.00044658183,0.9923375,0.00033441756,0.0050325026],"study_design_scores_gemma":[0.000009955657,0.000015760575,0.000082355364,0.000027341779,0.000010707388,0.00016663726,0.00006776887,0.013823203,0.0014326136,0.98040926,0.0039353822,0.000018998087],"about_ca_topic_score_codex":0.0007156033,"about_ca_topic_score_gemma":0.00050014135,"teacher_disagreement_score":0.0073026847,"about_ca_system_score_codex":0.0011419263,"about_ca_system_score_gemma":0.0012365475,"threshold_uncertainty_score":0.029464006},"labels":[],"label_agreement":null},{"id":"W2021026416","doi":"10.1145/2372251.2372270","title":"Evolution of features and their dependencies - an explorative study in OSS","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Provincia autonoma di Bolzano - Alto Adige","keywords":"Commit; Computer science; Feature (linguistics); Software engineering; Source code; Graph; Process (computing); Software maintenance; Software product line; Software; Data mining; Software development; Database; Programming language; Theoretical computer science","score_opus":0.03540925746221999,"score_gpt":0.29261447962481374,"score_spread":0.2572052221625937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021026416","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.970763,0.00025870572,0.027045423,0.000169565,0.0000034431696,0.000048382513,0.00011264977,0.00006847257,0.0015302561],"genre_scores_gemma":[0.9839028,0.00015504254,0.015301063,0.000019073457,0.0000037564744,0.000041711664,0.00013183021,0.000043859814,0.00040073818],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99772483,0.0012845587,0.00010499081,0.00030326168,0.0004838793,0.00009847851],"domain_scores_gemma":[0.97147155,0.022065025,0.0031442957,0.0020816622,0.000980628,0.00025684034],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029364354,0.0003464542,0.00027670845,0.0028264162,0.0006529665,0.0011311814,0.0006509033,0.00051301945,0.00045786312],"category_scores_gemma":[0.020171907,0.00044733484,0.000561228,0.0018371335,0.0014201049,0.002446015,0.0010875335,0.0007037585,0.00007703988],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003099468,0.0006017479,0.60109353,0.00038707824,0.00030594997,0.0038958301,0.039566025,0.0591666,0.040499233,0.02234874,0.000755217,0.23107015],"study_design_scores_gemma":[0.000035132674,0.0008290232,0.7245535,0.00024440957,0.00021761522,0.003278744,0.017573286,0.19020456,0.017293205,0.03357871,0.01200309,0.00018883043],"about_ca_topic_score_codex":0.0023727021,"about_ca_topic_score_gemma":0.0039386516,"teacher_disagreement_score":0.0029364354,"about_ca_system_score_codex":0.00068021426,"about_ca_system_score_gemma":0.00043177555,"threshold_uncertainty_score":0.015529573},"labels":[],"label_agreement":null},{"id":"W2021242474","doi":"10.1145/1453101.1453130","title":"Semi-automating small-scale source code reuse via structural correspondence","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Reuse; Context (archaeology); Code reuse; Source code; Unification; Jigsaw; Software engineering; Code (set theory); Scale (ratio); Usability; Database; Programming language; Software; Human–computer interaction; Engineering; Set (abstract data type)","score_opus":0.022142939863697173,"score_gpt":0.25229005088548556,"score_spread":0.2301471110217884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021242474","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.069277555,0.00007666044,0.92273533,0.00012596414,0.000010542343,0.00044565144,0.000049354196,0.0060297647,0.0012492127],"genre_scores_gemma":[0.25360614,0.000070444024,0.744486,0.00003539474,0.0000067388223,0.00022576984,0.00023187802,0.00048463725,0.0008529823],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98972404,0.0047239894,0.00070998276,0.0014956162,0.0029588314,0.00038746907],"domain_scores_gemma":[0.9526503,0.026733207,0.003903428,0.012248917,0.004038655,0.00042540007],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075905416,0.000961304,0.0010534133,0.0026410248,0.0011839083,0.0018931167,0.0024199663,0.0012140394,0.0021914826],"category_scores_gemma":[0.041992653,0.0008079851,0.0010715342,0.0017430834,0.0017556763,0.0028131614,0.003926716,0.0013531819,0.0010410091],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047085938,0.000753923,0.014489273,0.00086871145,0.0001470848,0.0005074285,0.004560637,0.048620783,0.08825197,0.021847503,0.002611246,0.81687057],"study_design_scores_gemma":[0.00025303187,0.00095673994,0.008534013,0.0002214472,0.00017848816,0.0013112725,0.0021243212,0.71502554,0.20220397,0.051804904,0.017214132,0.0001721691],"about_ca_topic_score_codex":0.0019960857,"about_ca_topic_score_gemma":0.003166107,"teacher_disagreement_score":0.0075905416,"about_ca_system_score_codex":0.00076336466,"about_ca_system_score_gemma":0.0035144861,"threshold_uncertainty_score":0.040143132},"labels":[],"label_agreement":null},{"id":"W2021330659","doi":"10.1109/icpc.2011.37","title":"Context and Vision: Studying Two Factors Impacting Program Comprehension","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Program comprehension; Comprehension; Computer science; Context (archaeology); Identifier; Cognition; Process (computing); Empirical research; Documentation; Cognitive psychology; Cognitive science; Human–computer interaction; Psychology; Programming language; Software system; Software; Epistemology","score_opus":0.08018862233186298,"score_gpt":0.3507900472435762,"score_spread":0.27060142491171324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021330659","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9917989,0.00024477617,0.0043453616,0.00014106308,0.0000063515404,0.00013196272,0.000054040018,0.00005690936,0.0032206255],"genre_scores_gemma":[0.9958872,0.00007195894,0.0035909319,0.000028969032,0.000004786361,0.00007534064,0.0000587477,0.000022941253,0.00025912124],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9936109,0.0036454727,0.000356732,0.0007429643,0.0012295699,0.00041425665],"domain_scores_gemma":[0.8363547,0.13126467,0.019870738,0.0029954207,0.0069636516,0.002550881],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047839163,0.00063029147,0.0004941793,0.002187373,0.0007277917,0.0026754974,0.0005700893,0.0009901631,0.0025850935],"category_scores_gemma":[0.09169262,0.0004307416,0.0005747651,0.0012010793,0.0012776435,0.0042994614,0.001709006,0.0011047865,0.00025711476],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002322882,0.0015902511,0.7687038,0.0014227836,0.000282654,0.00088074215,0.07827755,0.0025906307,0.029977698,0.0037833366,0.0007344443,0.109433115],"study_design_scores_gemma":[0.00012771186,0.0019844961,0.93258846,0.00032965493,0.00039997458,0.0005894465,0.026142437,0.014851012,0.013396417,0.006472705,0.0029345278,0.00018314544],"about_ca_topic_score_codex":0.0027070928,"about_ca_topic_score_gemma":0.0026789275,"teacher_disagreement_score":0.0047839163,"about_ca_system_score_codex":0.0011985381,"about_ca_system_score_gemma":0.0013857564,"threshold_uncertainty_score":0.025300026},"labels":[],"label_agreement":null},{"id":"W2021375786","doi":"10.1145/1809198.1809203","title":"The implications of how we tag software artifacts","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Metadata; Computer science; Semantics (computer science); Software; Source code; Simplicity; Software engineering; World Wide Web; Code (set theory); Information retrieval; Data science; Programming language; Set (abstract data type)","score_opus":0.019815370870302083,"score_gpt":0.26871372644817465,"score_spread":0.24889835557787257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021375786","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08130517,0.008318169,0.5002064,0.28048474,0.0029522153,0.00023514117,0.00048634733,0.00056273746,0.125449],"genre_scores_gemma":[0.8391392,0.0041672094,0.1306812,0.011116803,0.0010277161,0.00025390412,0.0003116739,0.00056696497,0.012735337],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.94999266,0.038814254,0.0019739731,0.0042965095,0.003592629,0.0013298591],"domain_scores_gemma":[0.8633379,0.09423657,0.0076567647,0.017135808,0.014846688,0.002786236],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029254597,0.0010974325,0.0005794388,0.0045731957,0.007363702,0.016154222,0.0028152983,0.006789215,0.0042728106],"category_scores_gemma":[0.118170716,0.0009957303,0.0008038683,0.0064610406,0.03620411,0.049015433,0.006228475,0.006735422,0.0019095904],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000096805416,0.00008367358,0.010745859,0.00027939168,0.00007683718,0.00040583423,0.02805306,0.0014391716,0.0012939334,0.88280064,0.009910789,0.064814106],"study_design_scores_gemma":[0.000024448596,0.000035063807,0.004141912,0.00024163937,0.000044511573,0.00053681736,0.019185185,0.0023360264,0.0015738633,0.9241865,0.047615543,0.000078436766],"about_ca_topic_score_codex":0.014726933,"about_ca_topic_score_gemma":0.013268505,"teacher_disagreement_score":0.029254597,"about_ca_system_score_codex":0.0055505326,"about_ca_system_score_gemma":0.0039962977,"threshold_uncertainty_score":0.15471494},"labels":[],"label_agreement":null},{"id":"W2021420643","doi":"10.1109/vissoft.2014.14","title":"Slicing-Based Techniques for Visualizing Large Metamodels","year":2014,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Metamodeling; Slicing; Computer science; Visualization; Program slicing; Domain (mathematical analysis); Focus (optics); Software engineering; Task (project management); Human–computer interaction; Artificial intelligence; Systems engineering; Engineering; World Wide Web","score_opus":0.03923877202282398,"score_gpt":0.34929779463680843,"score_spread":0.31005902261398444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021420643","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002855815,0.00028647907,0.99188644,0.00010459116,0.000018187997,0.000041977164,0.00015474299,0.0040666754,0.0005850285],"genre_scores_gemma":[0.044789698,0.00062595104,0.9514652,0.00005896922,0.000019528314,0.0001429501,0.00064196426,0.0017225192,0.00053315907],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991704,0.00029664292,0.000086152584,0.00012858678,0.0002630686,0.00005511221],"domain_scores_gemma":[0.9959557,0.00230418,0.0002852689,0.00090220035,0.0004153362,0.00013734949],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018372489,0.0016405751,0.00066838093,0.00219853,0.0006211077,0.0017283446,0.0010849651,0.001014513,0.0049098977],"category_scores_gemma":[0.005528914,0.0009248318,0.0014255057,0.0015075159,0.00091147755,0.0033014198,0.002530501,0.0021198844,0.0006708011],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004805765,0.00012681172,0.0024417283,0.001769262,0.00030253094,0.0009973866,0.007316845,0.06896926,0.19700451,0.15487254,0.021434573,0.544284],"study_design_scores_gemma":[0.00016335245,0.00016286045,0.0015858103,0.00045396938,0.0001867589,0.0012594191,0.00077850063,0.54478645,0.1536721,0.15333216,0.14340644,0.00021216783],"about_ca_topic_score_codex":0.0020812547,"about_ca_topic_score_gemma":0.0029436704,"teacher_disagreement_score":0.0049098977,"about_ca_system_score_codex":0.0006044642,"about_ca_system_score_gemma":0.0008618133,"threshold_uncertainty_score":0.016425252},"labels":[],"label_agreement":null},{"id":"W2021538299","doi":"10.1145/2377656.2377657","title":"Systematizing pragmatic software reuse","year":2012,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Waterloo","funders":"","keywords":"Computer science; Reuse; Variety (cybernetics); Software engineering; Task (project management); Plan (archaeology); Software development; Process (computing); Source code; Software; Metaphor; Human–computer interaction; Systems engineering; Programming language; Artificial intelligence; Engineering","score_opus":0.08412914294878511,"score_gpt":0.3257056353884473,"score_spread":0.24157649243966223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021538299","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048673127,0.0002140773,0.916039,0.0017796209,0.0000596262,0.0003479256,0.000070869886,0.0016899658,0.031125855],"genre_scores_gemma":[0.40849477,0.00026152856,0.5819048,0.00037916395,0.000030915457,0.00048192302,0.00023709638,0.00039579632,0.007813944],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.991389,0.005046333,0.0005200965,0.0011389261,0.0014474333,0.000458278],"domain_scores_gemma":[0.98701966,0.006278205,0.00085805764,0.004549233,0.0010385633,0.00025624002],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073005245,0.0010854407,0.00041287267,0.0016389132,0.0018834046,0.0038660457,0.0018199492,0.0017105897,0.005763175],"category_scores_gemma":[0.018234458,0.0009876089,0.0014627534,0.00085321354,0.011554603,0.007571033,0.007905814,0.0027621423,0.0008484352],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048351758,0.00009994249,0.0024374404,0.0003479735,0.000040245704,0.00026768987,0.011081992,0.013429065,0.0075278096,0.88389677,0.0023643079,0.078458324],"study_design_scores_gemma":[0.00011583921,0.0002474734,0.0015770944,0.00021131962,0.00009070793,0.00083305786,0.003397471,0.10894268,0.011024316,0.760729,0.11273045,0.0001006466],"about_ca_topic_score_codex":0.0045564556,"about_ca_topic_score_gemma":0.0052221986,"teacher_disagreement_score":0.0073005245,"about_ca_system_score_codex":0.0026245895,"about_ca_system_score_gemma":0.0042872247,"threshold_uncertainty_score":0.038609326},"labels":[],"label_agreement":null},{"id":"W2021546454","doi":"10.1109/vlhcc.2012.6344497","title":"Automatically locating relevant programming help online","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Context (archaeology); World Wide Web; The Internet; Software; Software engineering; Web search query; Information retrieval; Data science; Human–computer interaction; Search engine; Programming language","score_opus":0.025241990829763618,"score_gpt":0.29542077727566024,"score_spread":0.2701787864458966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021546454","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6910527,0.003145742,0.24074487,0.0015320351,0.00020980525,0.0018373535,0.0058423784,0.034156233,0.02147891],"genre_scores_gemma":[0.59333324,0.0010319196,0.38991562,0.00038423925,0.00011484547,0.0006767653,0.0064952397,0.001527274,0.0065208413],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950879,0.002153371,0.00031566422,0.00082024466,0.0013316222,0.00029112038],"domain_scores_gemma":[0.96084565,0.031250764,0.0021474925,0.0015920639,0.0033980878,0.0007658399],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033827415,0.0012549883,0.0016033326,0.008250673,0.0010880642,0.0026351367,0.0014004924,0.0015656413,0.006177704],"category_scores_gemma":[0.024105676,0.00047839005,0.00055309775,0.0032442652,0.00036783126,0.0033959781,0.0020482927,0.0005624509,0.0036784231],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027608778,0.001954162,0.0259944,0.003805764,0.00017293895,0.0011471722,0.007832684,0.0029059062,0.10597791,0.0043304563,0.03549489,0.80762285],"study_design_scores_gemma":[0.0015087544,0.0041017635,0.079924814,0.0013110836,0.0011925399,0.0038404835,0.028551782,0.27474654,0.38394988,0.02836053,0.19189018,0.00062167726],"about_ca_topic_score_codex":0.0016566615,"about_ca_topic_score_gemma":0.0030885364,"teacher_disagreement_score":0.008250673,"about_ca_system_score_codex":0.00067328097,"about_ca_system_score_gemma":0.0017652443,"threshold_uncertainty_score":0.02066642},"labels":[],"label_agreement":null},{"id":"W2021655726","doi":"10.1145/1985404.1985413","title":"Scalable clone detection using description logic","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Scalability; Cloud computing; Semantic reasoner; clone (Java method); Software; Semantic Web; Source code; Data mining; Software engineering; Database; Programming language; Information retrieval; Artificial intelligence; Operating system","score_opus":0.1031951321264783,"score_gpt":0.26766130597314614,"score_spread":0.16446617384666784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021655726","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02877709,0.00024703087,0.9618234,0.0001989391,0.00001578163,0.00012763005,0.00026761557,0.007770945,0.00077151635],"genre_scores_gemma":[0.2713463,0.0002280853,0.7251853,0.00014637076,0.000018951027,0.00013794813,0.0014396682,0.00039042035,0.0011070978],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964295,0.00073682837,0.00029400527,0.00071453094,0.0015930118,0.00023221782],"domain_scores_gemma":[0.9904946,0.0053399852,0.00097154814,0.0015944781,0.00142407,0.0001752123],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023832992,0.00086860976,0.0011468399,0.0040570116,0.0007492415,0.00256014,0.002092673,0.0011906679,0.001050581],"category_scores_gemma":[0.0107892435,0.0006709031,0.0017368706,0.0024920104,0.0012178429,0.004029522,0.0024290655,0.0012014655,0.00046835485],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004409117,0.0002429882,0.013515381,0.0005941411,0.00026921337,0.00079859514,0.00068097666,0.14867853,0.031558357,0.036158826,0.0057718996,0.7612902],"study_design_scores_gemma":[0.00005174556,0.00005435625,0.000942295,0.000028084647,0.00006738641,0.00030050726,0.00013214056,0.9378172,0.019703021,0.03686343,0.004001385,0.000038474136],"about_ca_topic_score_codex":0.008686191,"about_ca_topic_score_gemma":0.0065523414,"teacher_disagreement_score":0.008686191,"about_ca_system_score_codex":0.0020421874,"about_ca_system_score_gemma":0.0021552674,"threshold_uncertainty_score":0.01727128},"labels":[],"label_agreement":null},{"id":"W2022090474","doi":"10.1109/msr.2013.6624012","title":"Making sense of online code snippets","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Snippet; Android (operating system); Information retrieval; Source code; Code (set theory); World Wide Web; Callback; Programming language; Set (abstract data type); Operating system","score_opus":0.045255063716788044,"score_gpt":0.31583377834754506,"score_spread":0.270578714630757,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022090474","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26100734,0.0033452462,0.53762937,0.0053151622,0.0029632126,0.0015333996,0.06659902,0.07619252,0.045414764],"genre_scores_gemma":[0.4082851,0.0023925833,0.4901273,0.001444495,0.0010487279,0.0013408716,0.061223004,0.01350513,0.020632707],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9969502,0.0005202902,0.0002933806,0.0005720819,0.0015058975,0.00015807374],"domain_scores_gemma":[0.97232914,0.018168047,0.0018816259,0.0026964718,0.004553183,0.0003715542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022298906,0.0018533266,0.00062867947,0.011130695,0.0011212849,0.0032371737,0.0011692613,0.0015753972,0.008596874],"category_scores_gemma":[0.03326978,0.00060071703,0.0007455386,0.0060475734,0.0010166697,0.005218576,0.0031108563,0.0012860242,0.0051598363],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009837609,0.00021443567,0.022006424,0.00231117,0.00028220654,0.007280876,0.008438331,0.005804588,0.030050104,0.014532982,0.12018518,0.78791],"study_design_scores_gemma":[0.00013439615,0.00034767276,0.05696535,0.0023308448,0.00042218782,0.00795738,0.011250207,0.13131055,0.0801463,0.092287436,0.6163371,0.0005106043],"about_ca_topic_score_codex":0.001888973,"about_ca_topic_score_gemma":0.003765158,"teacher_disagreement_score":0.011130695,"about_ca_system_score_codex":0.0004570965,"about_ca_system_score_gemma":0.0008634692,"threshold_uncertainty_score":0.02875936},"labels":[],"label_agreement":null},{"id":"W2022429257","doi":"10.1145/2591062.2591075","title":"DASHboards: enhancing developer situational awareness","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; BitTorrent tracker; Dash; Situation awareness; sort; Situational ethics; Data science; World Wide Web; Human–computer interaction; Software engineering; Engineering","score_opus":0.01798889216864847,"score_gpt":0.2713388981247568,"score_spread":0.25335000595610835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022429257","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1968188,0.002131999,0.6200701,0.0032874276,0.0011177467,0.00094380736,0.0027133932,0.15724126,0.015675426],"genre_scores_gemma":[0.56613135,0.0015947882,0.40784544,0.001016867,0.00039787905,0.0008701819,0.0053240047,0.0046047615,0.0122146],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961153,0.0014159252,0.00037228226,0.0008354372,0.0010708583,0.00019020303],"domain_scores_gemma":[0.96699256,0.017824426,0.0034057565,0.00641953,0.00298611,0.0023717063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058328724,0.0017099874,0.00066484226,0.002697351,0.0006612688,0.0028416773,0.002104257,0.001433457,0.0038990874],"category_scores_gemma":[0.038626272,0.00091660843,0.00046102385,0.001728922,0.00057466346,0.0076927124,0.0057572518,0.0023308452,0.0014850256],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014566662,0.0016136362,0.038679965,0.0014353291,0.00024983776,0.0009193757,0.008383079,0.008911053,0.03583999,0.010488316,0.06169281,0.8303299],"study_design_scores_gemma":[0.0014959676,0.0034863057,0.06283421,0.0009501544,0.0009984671,0.0020020246,0.007672002,0.29613486,0.0844473,0.06710769,0.4719348,0.0009362017],"about_ca_topic_score_codex":0.0018564343,"about_ca_topic_score_gemma":0.0028554082,"teacher_disagreement_score":0.0058328724,"about_ca_system_score_codex":0.00036909472,"about_ca_system_score_gemma":0.0012575537,"threshold_uncertainty_score":0.03084749},"labels":[],"label_agreement":null},{"id":"W2022461105","doi":"10.1007/s11334-005-0012-2","title":"Software release planning for evolving systems","year":2005,"lang":"en","type":"article","venue":"Innovations in Systems and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Set (abstract data type); Component (thermodynamics); Hierarchy; Process (computing); Weighting; Software release life cycle; Software; Analytic hierarchy process; Resource (disambiguation); Software system; Risk analysis (engineering); Operations research; Engineering","score_opus":0.021572828620047158,"score_gpt":0.26691134417449025,"score_spread":0.24533851555444308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022461105","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1279216,0.0012842872,0.8605995,0.00044034095,0.00009227768,0.0002791769,0.00012896144,0.0012002378,0.008053548],"genre_scores_gemma":[0.7806473,0.0006917572,0.21205813,0.00005392934,0.00004815458,0.00014240488,0.00029869363,0.00038196286,0.0056776553],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99848735,0.00059471594,0.00010763569,0.00016689108,0.0004459521,0.00019739282],"domain_scores_gemma":[0.99480224,0.003406458,0.0003734464,0.00045084357,0.0006663011,0.00030063317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033052722,0.00066576037,0.0006808904,0.0009642692,0.00066980516,0.0012754246,0.0011431332,0.0006569082,0.0030525576],"category_scores_gemma":[0.010405554,0.00082490296,0.0007513118,0.0006195348,0.00075322716,0.0014488961,0.00096262054,0.0013330906,0.0005122725],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047983896,0.00011889437,0.0015034016,0.00033138896,0.0000576041,0.00044533313,0.00042583537,0.7742171,0.006215757,0.024576038,0.0037560712,0.18787257],"study_design_scores_gemma":[0.000055134675,0.00018115414,0.00060482696,0.000034495522,0.000040879953,0.00009322402,0.00012478785,0.97695255,0.0031169339,0.01693785,0.0018389728,0.000019242616],"about_ca_topic_score_codex":0.004117589,"about_ca_topic_score_gemma":0.0048162206,"teacher_disagreement_score":0.004117589,"about_ca_system_score_codex":0.00072258705,"about_ca_system_score_gemma":0.0011287616,"threshold_uncertainty_score":0.017480135},"labels":[],"label_agreement":null},{"id":"W2022628966","doi":"10.1007/s11219-006-9010-3","title":"Measuring size, complexity, and coupling of hypergraph abstractions of software: An information-theory approach","year":2007,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Waterloo; Mississippi State University; National Science Foundation","keywords":"Computer science; Software metric; Software sizing; Theoretical computer science; Hypergraph; Software; Data mining; Software development; Software quality; Software construction; Programming language; Mathematics","score_opus":0.07791100068699204,"score_gpt":0.32054010421443907,"score_spread":0.24262910352744704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022628966","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6275755,0.00031038464,0.36911482,0.00030108308,0.000013731207,0.00009003819,0.00018747164,0.00032652568,0.0020804028],"genre_scores_gemma":[0.95623696,0.0000880609,0.04324974,0.00001765961,0.000019343897,0.000047075137,0.00014906643,0.000034001503,0.00015802437],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948606,0.0017192672,0.0003437964,0.00068384653,0.0021278297,0.0002645543],"domain_scores_gemma":[0.87652314,0.097189665,0.010913453,0.008965368,0.0046238545,0.0017845267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050661764,0.00071063376,0.0011799363,0.0084818825,0.00087642757,0.0032324907,0.0017661217,0.0016091027,0.00082130724],"category_scores_gemma":[0.062482525,0.0008340247,0.0012097256,0.004522404,0.0027013996,0.012154467,0.002491046,0.0015565953,0.000093587005],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012468698,0.0009309625,0.19986613,0.0007465841,0.0010651087,0.00030790854,0.0039941045,0.38445655,0.03541561,0.14655481,0.0010567665,0.22435863],"study_design_scores_gemma":[0.00004581951,0.0003265275,0.043197405,0.000042159532,0.000298363,0.00016807651,0.0006671237,0.79965013,0.0111824535,0.1438069,0.00050298916,0.0001120323],"about_ca_topic_score_codex":0.0027625312,"about_ca_topic_score_gemma":0.0022770728,"teacher_disagreement_score":0.0084818825,"about_ca_system_score_codex":0.0018766131,"about_ca_system_score_gemma":0.0010297376,"threshold_uncertainty_score":0.026792824},"labels":[],"label_agreement":null},{"id":"W2022675432","doi":"10.1002/smr.539","title":"TIDIER: an identifier splitting approach using speech recognition techniques","year":2011,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Identifier; Computer science; Program comprehension; Source code; Maintainability; Unique identifier; Set (abstract data type); Documentation; Software; Code (set theory); Relation (database); Natural language processing; Comprehension; Information retrieval; Artificial intelligence; Data mining; Software engineering; Programming language; Software system","score_opus":0.06451466082664108,"score_gpt":0.2980101579526624,"score_spread":0.2334954971260213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022675432","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.074390486,0.0003442178,0.9126622,0.00019919586,0.0001798267,0.00025210273,0.00081705116,0.008545562,0.002609313],"genre_scores_gemma":[0.16712162,0.00018077002,0.8252863,0.00011127469,0.00006082331,0.00023457047,0.0021810392,0.0004921328,0.004331487],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985335,0.00029330057,0.00013495449,0.00057687674,0.00038487508,0.00007648124],"domain_scores_gemma":[0.9977958,0.0008138782,0.00030342847,0.0003430883,0.0006404937,0.000103273764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001243487,0.00096402806,0.0009395642,0.0028902406,0.0004815313,0.0012501243,0.0009074048,0.0007089279,0.0042892057],"category_scores_gemma":[0.0038902173,0.00029802124,0.000742906,0.0014592035,0.00046030525,0.0016133446,0.001484514,0.00089263474,0.003946458],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005517608,0.00012504154,0.003182843,0.00029640042,0.0000762169,0.0002979042,0.00075155904,0.0039574215,0.10815739,0.0019759764,0.0038136137,0.8768139],"study_design_scores_gemma":[0.00023088552,0.0008379838,0.013606615,0.000115935,0.00038839382,0.0022097493,0.0017906745,0.65014243,0.2669902,0.011116701,0.052356713,0.00021366132],"about_ca_topic_score_codex":0.0013289353,"about_ca_topic_score_gemma":0.0014168331,"teacher_disagreement_score":0.0042892057,"about_ca_system_score_codex":0.00036248294,"about_ca_system_score_gemma":0.0007629472,"threshold_uncertainty_score":0.014348865},"labels":[],"label_agreement":null},{"id":"W2022794931","doi":"10.1145/2652524.2652547","title":"A qualitative analysis of software build system changes and build ownership styles","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Deliverable; Computer science; Context (archaeology); Software development; Software; Focus (optics); Software engineering; Overhead (engineering); Software construction; Software system; Software peer review; Systems engineering; Operating system; Engineering","score_opus":0.035426478277793515,"score_gpt":0.32726746530412404,"score_spread":0.29184098702633055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022794931","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93758386,0.00016234978,0.023181152,0.0013424088,0.00002385503,0.00032063265,0.0016036297,0.00006748211,0.035714623],"genre_scores_gemma":[0.99338186,0.00007691442,0.0035046574,0.00010816116,0.0000030852677,0.00022304199,0.00024454543,0.000015971827,0.0024417853],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9970457,0.0015249364,0.000102855876,0.00021459443,0.0008212019,0.00029073303],"domain_scores_gemma":[0.95961314,0.03132886,0.002444311,0.00094125204,0.0046641976,0.0010082859],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.003556806,0.00014868869,0.0001483382,0.0022193433,0.0014053168,0.001974492,0.00049924396,0.0004926705,0.0055690086],"category_scores_gemma":[0.017664742,0.00018546893,0.0001647398,0.002482063,0.0030530235,0.0022916365,0.001338005,0.0008094566,0.00042482183],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005081562,0.00046259983,0.14709583,0.0011497354,0.000028523413,0.0013671609,0.6282281,0.0023974532,0.015344937,0.098815724,0.0057936655,0.09880804],"study_design_scores_gemma":[0.000042906908,0.00032052645,0.21183996,0.0005217853,0.000034772267,0.0007364642,0.70642173,0.006084168,0.0094220005,0.021938242,0.04255415,0.00008334181],"about_ca_topic_score_codex":0.007887259,"about_ca_topic_score_gemma":0.010468695,"teacher_disagreement_score":0.9964432,"about_ca_system_score_codex":0.0044356096,"about_ca_system_score_gemma":0.0024930858,"threshold_uncertainty_score":0.032182753},"labels":[],"label_agreement":null},{"id":"W2023019272","doi":"10.1109/scam.2011.23","title":"Recovering a Balanced Overview of Topics in a Software Domain","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Domain engineering; Domain analysis; Domain (mathematical analysis); Computer science; Feature-oriented domain analysis; Granularity; Domain model; Reuse; Identification (biology); Software engineering; Source code; Software; Code (set theory); Business domain; Data mining; Software development; Data science; Software construction; Programming language; Domain knowledge; Engineering; Business rule","score_opus":0.05962982282458056,"score_gpt":0.282705894301885,"score_spread":0.2230760714773044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2023019272","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4330145,0.006646874,0.5430897,0.0014462797,0.00014369139,0.0006023379,0.0057858196,0.0024929203,0.0067779077],"genre_scores_gemma":[0.58628607,0.0049799653,0.39010492,0.00020675009,0.00026306804,0.0009824189,0.013409409,0.00082287093,0.0029444655],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965125,0.0009559837,0.00041109917,0.0009911625,0.00088125555,0.00024803932],"domain_scores_gemma":[0.9853167,0.008215922,0.0015284425,0.0013828597,0.0029579096,0.0005981153],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048328005,0.0015084652,0.0014260508,0.017943794,0.001549425,0.0045682644,0.0010362748,0.0015753958,0.0016160628],"category_scores_gemma":[0.02124201,0.0013092341,0.0013914838,0.010690682,0.00065652677,0.0094604865,0.0033711018,0.0020085988,0.0012354197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010369763,0.0004857303,0.11610035,0.0031613812,0.0006231128,0.0017651645,0.033728242,0.016231893,0.062535554,0.017692491,0.020914633,0.72572464],"study_design_scores_gemma":[0.0002740203,0.0009044487,0.24769345,0.0018714912,0.0017687012,0.0046345457,0.053674154,0.30215907,0.037806433,0.17275105,0.17580324,0.00065936224],"about_ca_topic_score_codex":0.004735717,"about_ca_topic_score_gemma":0.0050234064,"teacher_disagreement_score":0.017943794,"about_ca_system_score_codex":0.0012896169,"about_ca_system_score_gemma":0.0019212095,"threshold_uncertainty_score":0.02555859},"labels":[],"label_agreement":null},{"id":"W2023617068","doi":"10.1007/s10664-009-9125-9","title":"An empirical study on the efficiency of different design pattern representations in UML class diagrams","year":2010,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal; Computer Research Institute of Montréal","funders":"","keywords":"Computer science; Applications of UML; Unified Modeling Language; Class diagram; UML tool; Notation; Program comprehension; Class (philosophy); Software design pattern; Visualization; Software engineering; Programming language; Software; Data mining; Artificial intelligence; Software system","score_opus":0.04501694003319926,"score_gpt":0.33526054527152854,"score_spread":0.29024360523832926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2023617068","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98876816,0.00029305616,0.009249814,0.00015017917,0.000008245389,0.00008897389,0.00016427097,0.00010791076,0.0011694058],"genre_scores_gemma":[0.9739378,0.00024499834,0.024437761,0.00004540674,0.000009889386,0.000082010636,0.0005194106,0.00011880019,0.00060401316],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.96737176,0.0220432,0.0034930452,0.0020838184,0.0044634542,0.00054472155],"domain_scores_gemma":[0.26298916,0.6911853,0.015875394,0.017589066,0.011536772,0.00082426314],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020279475,0.00059135346,0.0005523052,0.004107081,0.0006377702,0.0024850126,0.0012688441,0.0015012102,0.0018495983],"category_scores_gemma":[0.29988626,0.00049330207,0.00086996524,0.004410884,0.0011293422,0.0061036455,0.0012342698,0.0013123576,0.00041965992],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004695905,0.00604865,0.33100826,0.0020225325,0.00059599313,0.00045678357,0.02157659,0.019780736,0.023294758,0.0073039704,0.0017128526,0.58150285],"study_design_scores_gemma":[0.0014574697,0.010052675,0.5982939,0.0009289977,0.0019796195,0.0026957996,0.022569787,0.28812686,0.048190173,0.011944801,0.013413497,0.00034648934],"about_ca_topic_score_codex":0.0020949466,"about_ca_topic_score_gemma":0.0027990923,"teacher_disagreement_score":0.020279475,"about_ca_system_score_codex":0.0013698824,"about_ca_system_score_gemma":0.0010306141,"threshold_uncertainty_score":0.10724944},"labels":[],"label_agreement":null},{"id":"W2023925487","doi":"10.1109/ase.2013.6693113","title":"AutoComment: Mining question and answer sites for automatic comment generation","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":256,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Java; Programming language; Android (operating system); Maintainability; Codebase; Source code; Leverage (statistics); Static program analysis; Code (set theory); Software; Software engineering; Information retrieval; Software development; Artificial intelligence; Operating system","score_opus":0.032209821485548525,"score_gpt":0.28815731158042657,"score_spread":0.25594749009487805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2023925487","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.163908,0.0007018864,0.6895697,0.0017236758,0.00043326116,0.0037199736,0.021099973,0.11110929,0.0077342335],"genre_scores_gemma":[0.22257948,0.00021877151,0.72494805,0.0004112271,0.0002569504,0.0023923921,0.034424033,0.0039999476,0.010769087],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9920803,0.003300228,0.00045062046,0.0015709207,0.0022535946,0.0003444125],"domain_scores_gemma":[0.9220059,0.047599316,0.0057532643,0.007040182,0.016129129,0.0014722267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069541507,0.0023175252,0.0009974213,0.007654609,0.001118126,0.0016847253,0.00249424,0.0018485596,0.006646995],"category_scores_gemma":[0.04423887,0.0007027368,0.0010359061,0.0028754305,0.00079552137,0.00394147,0.003152924,0.0016539382,0.006432934],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094290037,0.0010284577,0.03275954,0.0030281313,0.00017334612,0.0008745971,0.0073359343,0.0044160304,0.06455786,0.0039660614,0.08819129,0.7927258],"study_design_scores_gemma":[0.0004925992,0.0012733423,0.04992307,0.0006732314,0.00025532336,0.0016473922,0.007530065,0.58535504,0.16501409,0.014739643,0.17268631,0.00040997096],"about_ca_topic_score_codex":0.0029173836,"about_ca_topic_score_gemma":0.005287884,"teacher_disagreement_score":0.007654609,"about_ca_system_score_codex":0.00094934605,"about_ca_system_score_gemma":0.0023841201,"threshold_uncertainty_score":0.036777496},"labels":[],"label_agreement":null},{"id":"W202396775","doi":"10.1016/j.scico.2006.02.001","title":"Introduction to the special issue on software analysis, evolution and reengineering","year":2006,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Business process reengineering; Computer science; Software engineering; Reverse engineering; Software evolution; Software system; Software; Software development; Component (thermodynamics); Legacy system; Component-based software engineering; Software construction; Systems engineering; Programming language; Engineering; Manufacturing engineering","score_opus":0.005930556694977833,"score_gpt":0.23681208914462343,"score_spread":0.2308815324496456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W202396775","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00095203455,0.12620676,0.037203647,0.041874576,0.72561556,0.00012595298,0.00090471393,0.000859711,0.06625705],"genre_scores_gemma":[0.003863512,0.063701995,0.009282454,0.01493515,0.7068909,0.0001257342,0.001363802,0.0009346174,0.19890183],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99885404,0.00019555974,0.0001293698,0.0002746433,0.00044971643,0.00009664711],"domain_scores_gemma":[0.9950023,0.0023683757,0.00026116698,0.00042702293,0.0011691097,0.0007720644],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017301356,0.0019457919,0.0027583148,0.0038174395,0.0011839202,0.0053738933,0.0017461375,0.0028999613,0.07180207],"category_scores_gemma":[0.0050655394,0.0007441874,0.0018597699,0.0029581361,0.0013702967,0.00578418,0.0020347135,0.0056547634,0.03086131],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028609049,0.00005718257,0.00014543736,0.00037813952,0.000024269375,0.00007854479,0.00004441231,0.00021363437,0.00048491976,0.007521261,0.917417,0.07360663],"study_design_scores_gemma":[0.00000878366,0.00004380807,0.00044030734,0.00016332667,0.000018688948,0.0002259523,0.000030734154,0.0003005298,0.00012575109,0.010677639,0.987949,0.000015432366],"about_ca_topic_score_codex":0.0005438083,"about_ca_topic_score_gemma":0.0015962814,"teacher_disagreement_score":0.07180207,"about_ca_system_score_codex":0.00094031467,"about_ca_system_score_gemma":0.0011812055,"threshold_uncertainty_score":0.24020189},"labels":[],"label_agreement":null},{"id":"W2024209326","doi":"10.1016/j.ins.2005.12.002","title":"Identification of defect-prone classes in telecommunication software systems using design metrics","year":2005,"lang":"en","type":"article","venue":"Information Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":99,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Negative binomial distribution; Computer science; Suite; Poisson regression; Software; Regression analysis; Data mining; Identification (biology); Poisson distribution; Statistics; Machine learning; Mathematics; Programming language; Population","score_opus":0.07743519655449478,"score_gpt":0.32720021633184165,"score_spread":0.24976501977734689,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2024209326","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8902373,0.00032145937,0.10738549,0.00012944143,0.000014126437,0.00010295365,0.00026032334,0.000769123,0.0007798095],"genre_scores_gemma":[0.95458984,0.00008129762,0.044572353,0.00001135482,0.000009506,0.000063802836,0.0003738514,0.00006276995,0.00023522206],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99686396,0.00081533974,0.00041745233,0.00037824613,0.001345412,0.00017960102],"domain_scores_gemma":[0.9501494,0.025415814,0.011920098,0.0027901349,0.008685101,0.0010394966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030472758,0.00085844286,0.00080287614,0.00862263,0.00046534676,0.0012360788,0.000992754,0.00092420285,0.00050921424],"category_scores_gemma":[0.028301325,0.0004122305,0.0007045084,0.0024626146,0.0005632889,0.0022463815,0.00081694865,0.000652898,0.00010969158],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005914911,0.00049879466,0.537718,0.0005643025,0.0003456121,0.0005071963,0.0012531555,0.10634142,0.036476836,0.0060902317,0.0012378102,0.30837515],"study_design_scores_gemma":[0.000050237348,0.00081421837,0.15487638,0.00009261835,0.00024523417,0.0006849662,0.00036274764,0.8109054,0.019731447,0.011055116,0.0011110759,0.000070591486],"about_ca_topic_score_codex":0.0026177382,"about_ca_topic_score_gemma":0.003811344,"teacher_disagreement_score":0.00862263,"about_ca_system_score_codex":0.00081449,"about_ca_system_score_gemma":0.0010418658,"threshold_uncertainty_score":0.016115725},"labels":[],"label_agreement":null},{"id":"W2024379951","doi":"10.1109/wcre.2010.18","title":"Extracting Sequence Diagrams from Execution Traces Using Interactive Visualization","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Fonds Québécois de la Recherche sur la Nature et les Technologies; Ministère du Développement Économique, de l’Innovation et de l’Exportation","keywords":"Sequence diagram; Computer science; Visualization; Sequence (biology); Unified Modeling Language; Set (abstract data type); Programming language; Data mining; Theoretical computer science; Software","score_opus":0.04281087544276049,"score_gpt":0.3487179159585597,"score_spread":0.3059070405157992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2024379951","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006589029,0.00009135353,0.9796245,0.000114964896,0.000017499693,0.0001272077,0.0004620547,0.012263355,0.00070996245],"genre_scores_gemma":[0.04908241,0.00019593685,0.947021,0.00003474349,0.000018994971,0.00021248391,0.0016869524,0.0010416112,0.0007058374],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99750155,0.0008898428,0.00026514076,0.00035549304,0.00089426915,0.00009371748],"domain_scores_gemma":[0.9868492,0.00812794,0.0012071937,0.0019775925,0.0016397485,0.00019819623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023842165,0.002200955,0.0009866232,0.00543566,0.00065754616,0.002621031,0.0013421688,0.0011548877,0.00443926],"category_scores_gemma":[0.015455918,0.0006938891,0.0011510471,0.0023719277,0.0006017691,0.0024075382,0.0019139224,0.0014715715,0.0015097182],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004435128,0.0002505183,0.004684076,0.001296154,0.0002021393,0.001100889,0.003512613,0.048001338,0.09750146,0.026365886,0.009922891,0.8067185],"study_design_scores_gemma":[0.00017909938,0.00033812542,0.005382434,0.0005219018,0.00019384328,0.001736982,0.0010073192,0.6453952,0.1523736,0.08384134,0.10872413,0.0003060383],"about_ca_topic_score_codex":0.0025863785,"about_ca_topic_score_gemma":0.0034122502,"teacher_disagreement_score":0.00543566,"about_ca_system_score_codex":0.0005308229,"about_ca_system_score_gemma":0.0017390994,"threshold_uncertainty_score":0.014850795},"labels":[],"label_agreement":null},{"id":"W2024402488","doi":"10.1016/j.infsof.2011.10.007","title":"Comparing alternatives for analyzing requirements trade-offs – In the absence of numerical data","year":2011,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Swap (finance); Computer science; Process (computing); Work (physics); Risk analysis (engineering); Operations research; Management science; Engineering; Economics","score_opus":0.0907674307304828,"score_gpt":0.3177942355904184,"score_spread":0.22702680485993562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2024402488","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6347637,0.0019726832,0.342165,0.001282163,0.00014952788,0.0005356219,0.0009894775,0.00047072582,0.017671118],"genre_scores_gemma":[0.9031327,0.0002620673,0.09493526,0.000103566184,0.000035494224,0.00026732532,0.0006138666,0.00007608243,0.00057363336],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96302277,0.025019294,0.0017207786,0.0019632073,0.0069747493,0.0012992594],"domain_scores_gemma":[0.6891217,0.29405943,0.005511441,0.0064800424,0.0036464373,0.0011808508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.036403604,0.0015990813,0.001918,0.008199806,0.000656793,0.0057849507,0.002306457,0.0031709655,0.005747085],"category_scores_gemma":[0.1577349,0.0008341769,0.0026493603,0.0048039313,0.0027949228,0.0062506,0.0022395977,0.0024937773,0.00067669386],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013839962,0.0017743862,0.0529799,0.0022888822,0.0021451476,0.00047008882,0.001319783,0.49088627,0.0052177627,0.14061293,0.003283521,0.2851814],"study_design_scores_gemma":[0.00037345345,0.0016105704,0.0107811345,0.00019123954,0.00027501903,0.00014267678,0.0010056696,0.86522925,0.0017030158,0.11712286,0.0013988141,0.00016628954],"about_ca_topic_score_codex":0.0019202837,"about_ca_topic_score_gemma":0.0029616442,"teacher_disagreement_score":0.036403604,"about_ca_system_score_codex":0.0025584265,"about_ca_system_score_gemma":0.001673361,"threshold_uncertainty_score":0.192523},"labels":[],"label_agreement":null},{"id":"W2024695825","doi":"10.1007/s10664-011-9163-y","title":"Qualitative research in software engineering","year":2011,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Qualitative research; Phenomenon; Ethnography; Context (archaeology); Social phenomenon; Social research; Participant observation; Grounded theory; Knowledge management; Software; Data science; Management science; Sociology; Qualitative property; Computer science; Epistemology; Engineering ethics; Social science; Engineering","score_opus":0.19842369192862866,"score_gpt":0.4223873846295447,"score_spread":0.22396369270091607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2024695825","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.217332,0.013626576,0.29470983,0.07199603,0.0015543542,0.0060259947,0.0018531522,0.00021609222,0.39268598],"genre_scores_gemma":[0.9140066,0.0033146208,0.043064762,0.0061399844,0.00009640473,0.0046408023,0.0003011167,0.00009680623,0.028338907],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9370529,0.055093925,0.0008564559,0.0011205202,0.0048263436,0.0010499202],"domain_scores_gemma":[0.8056062,0.17010708,0.0034622457,0.004528739,0.013665453,0.0026303732],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.048521698,0.00047567155,0.00063066615,0.0027378,0.0061306152,0.005181123,0.0018876087,0.0016606696,0.011945734],"category_scores_gemma":[0.09610756,0.00051718863,0.000357796,0.0028917568,0.013185265,0.0051966547,0.0043473183,0.002424273,0.0011797453],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012631238,0.0002979247,0.0046508987,0.0034532908,0.000026008627,0.00035260804,0.31227863,0.0006564006,0.0020137422,0.5894342,0.009687048,0.0770229],"study_design_scores_gemma":[0.00013778711,0.00027199576,0.0046967464,0.005490239,0.000032261458,0.0005447264,0.4536624,0.001308046,0.0037249885,0.3390623,0.1910088,0.000059747206],"about_ca_topic_score_codex":0.006340111,"about_ca_topic_score_gemma":0.008989803,"teacher_disagreement_score":0.9514783,"about_ca_system_score_codex":0.009027703,"about_ca_system_score_gemma":0.013176179,"threshold_uncertainty_score":0.25661033},"labels":[],"label_agreement":null},{"id":"W2025152749","doi":"10.1145/2486046.2486059","title":"A process practice to validate the quality of reused component documentation: a case study involving open-source components","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Reuse; Computer science; Software engineering; Process (computing); Component (thermodynamics); Quality (philosophy); Software quality; Software; Software documentation; Software development process; Software development; Engineering; Operating system","score_opus":0.09492014226052924,"score_gpt":0.4135468463924063,"score_spread":0.31862670413187705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025152749","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8187962,0.00046190582,0.16623382,0.0027459378,0.00005999187,0.0028972407,0.000056753568,0.0003337607,0.008414446],"genre_scores_gemma":[0.81879795,0.0003593274,0.17767613,0.00029965377,0.000016743737,0.00092429254,0.00006378957,0.00007821417,0.001783762],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9249723,0.052986797,0.0039585945,0.003918312,0.012527798,0.0016361948],"domain_scores_gemma":[0.82768774,0.096386634,0.012774638,0.027706949,0.032098856,0.0033451493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.065925166,0.0008524655,0.00057228794,0.003331336,0.004725567,0.0043078004,0.0030507804,0.0031770528,0.0009280063],"category_scores_gemma":[0.11532437,0.00071026204,0.00072085316,0.0023419098,0.004633123,0.005830702,0.005751419,0.0023665223,0.0004081902],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005429219,0.008551275,0.077914484,0.002141504,0.00015800125,0.006191287,0.33043218,0.008860971,0.027258504,0.018663539,0.0028955543,0.51638985],"study_design_scores_gemma":[0.001511373,0.023676103,0.13305403,0.007911982,0.0008686744,0.01625238,0.44058514,0.066564895,0.11472518,0.035455573,0.15833974,0.0010550189],"about_ca_topic_score_codex":0.0035344802,"about_ca_topic_score_gemma":0.005335064,"teacher_disagreement_score":0.065925166,"about_ca_system_score_codex":0.0044940216,"about_ca_system_score_gemma":0.009298431,"threshold_uncertainty_score":0.34864974},"labels":[],"label_agreement":null},{"id":"W2025486920","doi":"10.1145/1028976.1029002","title":"Recovering binary class relationships","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":85,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Programming language; Class (philosophy); Binary number; Theoretical computer science; Object-oriented programming; Software; Programming complexity; Software development; Software construction; Artificial intelligence; Mathematics","score_opus":0.040723446344669706,"score_gpt":0.2625467331104656,"score_spread":0.2218232867657959,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025486920","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08880476,0.00068298285,0.8934905,0.0010394728,0.00018989368,0.00011067403,0.0010140795,0.0029675898,0.011700037],"genre_scores_gemma":[0.4533394,0.00041885462,0.53400254,0.00029318893,0.00014968657,0.00018065241,0.002758544,0.0012582535,0.0075988267],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9936779,0.0010625462,0.00035814435,0.0012141453,0.0029588728,0.0007284243],"domain_scores_gemma":[0.98014945,0.0066718864,0.0021630581,0.00665762,0.0039062377,0.0004516144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031900713,0.0007990558,0.0009872844,0.0048975423,0.0020846466,0.004244287,0.0019952264,0.0026726706,0.0047222334],"category_scores_gemma":[0.04083381,0.00084063865,0.0009320004,0.0038288243,0.0015119683,0.008132106,0.0054439437,0.0027773755,0.0025791777],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002855199,0.00014324143,0.020921607,0.0003357263,0.000043758264,0.00072281057,0.0016639815,0.013107108,0.013871632,0.31799635,0.017374437,0.6135339],"study_design_scores_gemma":[0.00007169815,0.00006435131,0.0068428004,0.00015843117,0.00007464684,0.0011537105,0.001028569,0.28397915,0.025581691,0.56907225,0.11185871,0.00011395708],"about_ca_topic_score_codex":0.0053136894,"about_ca_topic_score_gemma":0.0037353437,"teacher_disagreement_score":0.0053136894,"about_ca_system_score_codex":0.001172914,"about_ca_system_score_gemma":0.002338134,"threshold_uncertainty_score":0.016870916},"labels":[],"label_agreement":null},{"id":"W2025748259","doi":"10.1016/j.bandc.2004.06.005","title":"The effects of uncertainty in error monitoring on associated ERPs","year":2004,"lang":"en","type":"article","venue":"Brain and Cognition","topic":"Software Engineering Research","field":"Computer Science","cited_by":205,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University; University of Waterloo","funders":"","keywords":"Psychology; Cognitive psychology; Audiology","score_opus":0.012615723068188225,"score_gpt":0.2661108959532865,"score_spread":0.2534951728850983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025748259","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9856431,0.0006170861,0.00917284,0.00020674775,0.00005735497,0.000042335327,0.00020328148,0.00006603626,0.0039913366],"genre_scores_gemma":[0.99616903,0.00021974974,0.002875624,0.000069654765,0.00008753462,0.00002603747,0.00010494491,0.00008082984,0.00036658082],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9989497,0.00027514543,0.000081778584,0.00022223237,0.00036387568,0.00010721232],"domain_scores_gemma":[0.9575409,0.03730404,0.0025817775,0.0010502401,0.00093803444,0.00058498775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020039387,0.0006020422,0.00043262285,0.00057019613,0.00031236763,0.0012650044,0.0004120403,0.0007928054,0.003119764],"category_scores_gemma":[0.058502194,0.0005378123,0.00034748693,0.0005629464,0.0004573772,0.0015032679,0.00072505063,0.001389837,0.00022021485],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.03296056,0.0012829176,0.10622776,0.0010432483,0.00059062324,0.0013275812,0.003507016,0.015330325,0.5993876,0.0076518236,0.0014115345,0.22927895],"study_design_scores_gemma":[0.00023151955,0.0015633829,0.9274328,0.00011271455,0.00043105547,0.0014004201,0.00033415618,0.021309407,0.032258056,0.01398239,0.0008345672,0.00010963422],"about_ca_topic_score_codex":0.0010753778,"about_ca_topic_score_gemma":0.0008721535,"teacher_disagreement_score":0.003119764,"about_ca_system_score_codex":0.00026113456,"about_ca_system_score_gemma":0.00030982323,"threshold_uncertainty_score":0.010597944},"labels":[],"label_agreement":null},{"id":"W2025893091","doi":"10.1109/icsm.2011.6080814","title":"Source code comprehension strategies and metrics to predict comprehension effort in software maintenance and evolution tasks - an empirical study with industry practitioners","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Program comprehension; Computer science; Comprehension; Code refactoring; Software maintenance; Source code; Semantics (computer science); Software metric; Programming language; Task (project management); Static program analysis; Empirical research; Software development; Artificial intelligence; Natural language processing; Software quality; Software; Software system; Statistics; Engineering","score_opus":0.045774602770144265,"score_gpt":0.30274283293018284,"score_spread":0.25696823016003856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025893091","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9985267,0.000053462158,0.001084373,0.000020663338,0.0000012197958,0.00003446079,0.00002568109,0.0000108555,0.00024256084],"genre_scores_gemma":[0.9978288,0.0000345955,0.0017399503,0.000015632037,0.0000037676082,0.00007476684,0.00012020864,0.000010772973,0.0001715152],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98889714,0.0059255133,0.0010964569,0.0014578012,0.0022620517,0.00036091372],"domain_scores_gemma":[0.71106446,0.22681123,0.030945443,0.0095014,0.019017562,0.0026599474],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0153463865,0.00063562,0.00043607075,0.0024441679,0.00032792083,0.0013543514,0.0007850051,0.0013287106,0.00089775334],"category_scores_gemma":[0.14816335,0.00044042594,0.00039491203,0.0013211928,0.0008175694,0.0027964185,0.0011091607,0.0010273609,0.000371025],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003425901,0.0013329295,0.95605075,0.0001427439,0.00014087129,0.00012277828,0.008816316,0.0012976348,0.0030981358,0.0000984707,0.00022218011,0.02833468],"study_design_scores_gemma":[0.000054519038,0.002065583,0.97557247,0.000041528936,0.00006647447,0.00033239328,0.004049763,0.014506985,0.0024421415,0.00029567475,0.00052986667,0.000042520904],"about_ca_topic_score_codex":0.0012601491,"about_ca_topic_score_gemma":0.0017406248,"teacher_disagreement_score":0.9846536,"about_ca_system_score_codex":0.00051102793,"about_ca_system_score_gemma":0.00038906123,"threshold_uncertainty_score":0.081160426},"labels":[],"label_agreement":null},{"id":"W2025962632","doi":"10.1109/icsme.2014.54","title":"Evaluating Modern Clone Detection Tools","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":122,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Benchmark (surveying); Computer science; Precision and recall; Cloning (programming); Software maintenance; Matching (statistics); Code (set theory); Software; Code refactoring; Software engineering; Machine learning; Artificial intelligence; Data mining; Software system; Programming language; Biology; Genetics","score_opus":0.07222274584709264,"score_gpt":0.3402158347236393,"score_spread":0.26799308887654666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025962632","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75657254,0.010585009,0.17440346,0.001296391,0.00049715204,0.0010791488,0.0047910307,0.03809683,0.012678416],"genre_scores_gemma":[0.62183934,0.0014771526,0.35384297,0.0005470419,0.00013390569,0.0005316274,0.016579371,0.001865758,0.003182827],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9307902,0.015386608,0.008024585,0.009583535,0.033941854,0.0022732106],"domain_scores_gemma":[0.8010332,0.10150406,0.018786509,0.023219392,0.05216711,0.0032896958],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.036002282,0.002141275,0.0014291758,0.01923999,0.0014712687,0.0047178003,0.0051220083,0.00363758,0.0012314521],"category_scores_gemma":[0.14625677,0.00085366983,0.0017791069,0.008186048,0.0014580792,0.006465477,0.0047542704,0.001563793,0.0009014683],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011571065,0.0010432128,0.17316464,0.0029771132,0.0011946985,0.0005742157,0.0064429897,0.024949972,0.021471126,0.0077977353,0.025206089,0.7340211],"study_design_scores_gemma":[0.00070747273,0.004665644,0.2328762,0.0022014068,0.0014576883,0.0038800947,0.005993622,0.49381724,0.092231505,0.013630029,0.14748396,0.0010551641],"about_ca_topic_score_codex":0.014887665,"about_ca_topic_score_gemma":0.016963134,"teacher_disagreement_score":0.9639977,"about_ca_system_score_codex":0.004461081,"about_ca_system_score_gemma":0.003973989,"threshold_uncertainty_score":0.19040054},"labels":[],"label_agreement":null},{"id":"W2026170849","doi":"10.1007/s10664-014-9338-4","title":"On rapid releases and software testing: a case study and a semi-systematic literature review","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":113,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Context (archaeology); Test suite; Systematic review; Scope (computer science); Software engineering; Software; Software bug; Software testing; Test case; Operating system","score_opus":0.027427958343474917,"score_gpt":0.2857589749436667,"score_spread":0.25833101660019175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2026170849","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46784762,0.4406667,0.024843963,0.010483641,0.00060904433,0.00710003,0.0029293322,0.00016670446,0.045353007],"genre_scores_gemma":[0.65670973,0.30823013,0.02434525,0.002848104,0.00023357927,0.002498043,0.0016678493,0.0000873259,0.0033799873],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9768734,0.011182297,0.0041964217,0.0011131867,0.0058882358,0.0007464995],"domain_scores_gemma":[0.7839958,0.18809272,0.011231318,0.003712899,0.01197508,0.000992185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023031792,0.00061496143,0.0011464757,0.015959714,0.0016662847,0.0024295892,0.0015555552,0.0017761767,0.002354048],"category_scores_gemma":[0.06054345,0.00053625624,0.0010352304,0.0151587995,0.0018755641,0.0037718443,0.0026426166,0.0011660063,0.0003995851],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004212466,0.0008285102,0.024407139,0.14233932,0.0006001983,0.015050514,0.08042626,0.0013757391,0.0072703687,0.011337614,0.0106318565,0.7053112],"study_design_scores_gemma":[0.00019318188,0.0020177546,0.094265185,0.36184445,0.0028993608,0.014482143,0.20012529,0.001517544,0.011379422,0.008958283,0.30197826,0.00033918786],"about_ca_topic_score_codex":0.004378614,"about_ca_topic_score_gemma":0.013707575,"teacher_disagreement_score":0.023031792,"about_ca_system_score_codex":0.003465379,"about_ca_system_score_gemma":0.015954584,"threshold_uncertainty_score":0.12180519},"labels":[],"label_agreement":null},{"id":"W2026316743","doi":"10.1145/1295074.1295082","title":"Toward a text classification system for the quality assessment of software requirements written in natural language","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software engineering; Documentation; Artifact (error); Natural language; Quality (philosophy); Process (computing); Software quality; Software requirements; Software requirements specification; Requirements engineering; Requirements elicitation; Software; Natural language processing; Artificial intelligence; Software development; Software design; Programming language","score_opus":0.09265065982244883,"score_gpt":0.3974940928867056,"score_spread":0.30484343306425676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2026316743","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02536261,0.00021330797,0.9461533,0.00056484836,0.000076863624,0.0009992897,0.0012901209,0.023736876,0.0016028871],"genre_scores_gemma":[0.080764234,0.00012879453,0.9120961,0.00015795023,0.00007655432,0.0010854359,0.0036587024,0.00038643708,0.0016458117],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916276,0.002715875,0.0015668105,0.0015366266,0.002317236,0.00023587326],"domain_scores_gemma":[0.95178384,0.021648582,0.006027549,0.0033370943,0.016422527,0.00078039814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009987628,0.001136275,0.0014711588,0.009982699,0.0013411875,0.0042850226,0.0018647185,0.0020303002,0.00265442],"category_scores_gemma":[0.036020722,0.00041632616,0.0011286417,0.0054356675,0.0008714056,0.0044098035,0.0011819184,0.0017911847,0.0031203763],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079840835,0.0005969255,0.013952969,0.00058391417,0.00012979849,0.00029833004,0.0010526861,0.007225314,0.034120128,0.007933034,0.012113443,0.9211952],"study_design_scores_gemma":[0.00024905373,0.0007260361,0.01564922,0.00032170865,0.00031484797,0.0006149645,0.0007337346,0.8933716,0.04703485,0.017234463,0.02356866,0.0001809442],"about_ca_topic_score_codex":0.0059936754,"about_ca_topic_score_gemma":0.0041903025,"teacher_disagreement_score":0.009987628,"about_ca_system_score_codex":0.0022063681,"about_ca_system_score_gemma":0.0023050737,"threshold_uncertainty_score":0.052820206},"labels":[],"label_agreement":null},{"id":"W2026802633","doi":"10.1145/568760.568769","title":"Using genetic algorithms and coupling measures to devise optimal integration test orders","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":97,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Coupling (piping); Class (philosophy); Genetic algorithm; Object (grammar); Algorithm; Machine learning; Artificial intelligence; Engineering","score_opus":0.06358410426365688,"score_gpt":0.29025348506535226,"score_spread":0.22666938080169538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2026802633","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055060074,0.00017750586,0.94191974,0.00015095892,0.000017942286,0.00013544122,0.000021011112,0.000526926,0.001990486],"genre_scores_gemma":[0.354415,0.00014213781,0.6440906,0.00010872287,0.000019211713,0.00022983074,0.00008726843,0.00018978812,0.000717442],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972254,0.0012081781,0.00018688703,0.0002749963,0.00085248763,0.00025212756],"domain_scores_gemma":[0.9898191,0.0074888268,0.0010032852,0.00053232064,0.00094999373,0.00020648986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003607428,0.0018310386,0.0014292424,0.0027654988,0.00056708243,0.0014795759,0.0011171359,0.001376408,0.0012125613],"category_scores_gemma":[0.01914762,0.0006968845,0.0006894022,0.0011071496,0.001332831,0.0014913908,0.0010143748,0.0015197651,0.00025289727],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006675181,0.00013178056,0.001583558,0.000065396875,0.000050253835,0.000060516726,0.000076676486,0.905561,0.00321839,0.009656451,0.00025469504,0.07927462],"study_design_scores_gemma":[0.00003855516,0.00012038882,0.00039221018,0.000018745684,0.000028244385,0.000022748849,0.000030773466,0.98561364,0.002593479,0.010767077,0.00035880707,0.000015368423],"about_ca_topic_score_codex":0.0041791354,"about_ca_topic_score_gemma":0.0053689457,"teacher_disagreement_score":0.0041791354,"about_ca_system_score_codex":0.0015932423,"about_ca_system_score_gemma":0.0038525206,"threshold_uncertainty_score":0.019078076},"labels":[],"label_agreement":null},{"id":"W2026940905","doi":"10.1016/j.infsof.2009.11.003","title":"Studying the impact of uncertainty in operational release planning – An integrated method and its initial evaluation","year":2009,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Staffing; Variance (accounting); Monte Carlo method; Baseline (sea); Process (computing); Computer science; Operations research; Uncertainty analysis; Heuristic; Reliability engineering; Industrial engineering; Sensitivity analysis; Span (engineering); Engineering; Simulation; Statistics; Mathematics; Civil engineering","score_opus":0.03909248299080244,"score_gpt":0.3792799754647623,"score_spread":0.34018749247395985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2026940905","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19836846,0.00096250477,0.79667634,0.00013356328,0.000040313793,0.0003577401,0.00014284781,0.0002912392,0.0030271253],"genre_scores_gemma":[0.6646545,0.00056641334,0.3331645,0.000032780878,0.000043657616,0.00032807633,0.00018209229,0.00011506747,0.00091291114],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9909943,0.0056233583,0.00034044945,0.00064594,0.002141366,0.00025467647],"domain_scores_gemma":[0.92956704,0.06211053,0.001788795,0.0026767314,0.0034263742,0.00043059702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01499936,0.0013468234,0.0014617696,0.0021281585,0.0007336844,0.0019593895,0.0019323501,0.0016214139,0.0022714983],"category_scores_gemma":[0.041365843,0.00085731153,0.0012994795,0.0026886426,0.0011324513,0.0041348515,0.0019695144,0.0018147504,0.00018946108],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013547621,0.000761123,0.011802032,0.0006280744,0.00047714313,0.00011457737,0.0006478578,0.7164479,0.005297481,0.013416103,0.00033900223,0.24871401],"study_design_scores_gemma":[0.00006939528,0.0008707406,0.0033842758,0.000036402074,0.000121731544,0.00004196135,0.00012615958,0.98771983,0.0029390105,0.0042860033,0.00036261592,0.000041840394],"about_ca_topic_score_codex":0.007957081,"about_ca_topic_score_gemma":0.004994674,"teacher_disagreement_score":0.01499936,"about_ca_system_score_codex":0.0016009124,"about_ca_system_score_gemma":0.0019223702,"threshold_uncertainty_score":0.07932514},"labels":[],"label_agreement":null},{"id":"W2027161770","doi":"10.5220/0005236100500061","title":"A Toolset for Simulink - Improving Software Engineering Practices in Development with Simulink","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Code refactoring; Software engineering; Modularity (biology); Program slicing; Slicing; Software; Embedded system; Programming language","score_opus":0.05789074993025396,"score_gpt":0.2982376652570027,"score_spread":0.24034691532674873,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2027161770","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004391275,0.000054137086,0.94386524,0.00009009189,0.000050955005,0.0001910792,0.00042095856,0.048458423,0.0024778084],"genre_scores_gemma":[0.045055717,0.00024193965,0.93932,0.000086255626,0.000019401012,0.00051692757,0.0017517594,0.009016047,0.003991917],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981402,0.0005319884,0.00029246122,0.000234921,0.00070307695,0.000097336444],"domain_scores_gemma":[0.99238974,0.0048796586,0.00046729375,0.0010459629,0.0010242572,0.00019312292],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003201336,0.0014446188,0.0004758956,0.0017394157,0.00043705892,0.0013048291,0.0016244078,0.00071775063,0.010458742],"category_scores_gemma":[0.011064592,0.0008947201,0.0007968077,0.0008583051,0.000719819,0.002035028,0.0020321761,0.0017840483,0.004441695],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007499885,0.00085114443,0.0046559833,0.0027528917,0.00020688532,0.0016392601,0.0024258865,0.14915244,0.1016515,0.05950843,0.039632738,0.63677275],"study_design_scores_gemma":[0.0005046967,0.00042680273,0.0015146586,0.0008826054,0.00013691795,0.001241988,0.00017171999,0.46379068,0.2179589,0.022722838,0.29045516,0.00019294978],"about_ca_topic_score_codex":0.0009829431,"about_ca_topic_score_gemma":0.0010762397,"teacher_disagreement_score":0.010458742,"about_ca_system_score_codex":0.0005000131,"about_ca_system_score_gemma":0.0012797577,"threshold_uncertainty_score":0.034987986},"labels":[],"label_agreement":null},{"id":"W2027204548","doi":"10.1145/2557833.2560584","title":"Leveraging machine learning and information retrieval techniques in software evolution tasks","year":2014,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software; Software engineering; Conjunction (astronomy); Software evolution; Artificial intelligence; Machine learning; Software development; Software construction; Programming language","score_opus":0.008723678462936355,"score_gpt":0.2292670223641095,"score_spread":0.22054334390117314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2027204548","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14685042,0.015168788,0.82272226,0.0033813266,0.00037600618,0.0007925079,0.0005263345,0.0054295803,0.004752765],"genre_scores_gemma":[0.4924073,0.0032558558,0.49782518,0.00089286023,0.00077438145,0.0003433311,0.001412659,0.00030747693,0.0027809876],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9899428,0.0054204613,0.0011741231,0.001056914,0.0019237181,0.00048206648],"domain_scores_gemma":[0.9664039,0.026221909,0.0014073078,0.0022330545,0.0033577583,0.00037608837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011860842,0.0016383063,0.0028691252,0.010649863,0.0010923232,0.003308469,0.002123253,0.0036722615,0.001511661],"category_scores_gemma":[0.04095045,0.0007395786,0.0016951583,0.009057451,0.0007671538,0.008096345,0.0022516737,0.0027038779,0.0015600979],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042090414,0.0008043737,0.005179395,0.00052625913,0.00033026683,0.00013786505,0.0004229926,0.047402,0.008590748,0.0025086089,0.006299777,0.9273768],"study_design_scores_gemma":[0.00012690306,0.00045104505,0.004201168,0.00006402086,0.00023717275,0.00030431114,0.00029671626,0.9573714,0.011570789,0.020856863,0.004403175,0.00011636876],"about_ca_topic_score_codex":0.0053704204,"about_ca_topic_score_gemma":0.005348799,"teacher_disagreement_score":0.011860842,"about_ca_system_score_codex":0.0011898313,"about_ca_system_score_gemma":0.0015863146,"threshold_uncertainty_score":0.062726915},"labels":[],"label_agreement":null},{"id":"W2027370816","doi":"10.1109/icpc.2010.48","title":"Understanding and Auditing the Licensing of Open Source Software Distributions","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":71,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"License; Computer science; Source code; Software; Open source; Software engineering; MIT License; Audit; Open source software; Operating system; Compatibility (geochemistry); Database; Engineering; Accounting","score_opus":0.07512046070146802,"score_gpt":0.3014063277299697,"score_spread":0.2262858670285017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2027370816","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9644319,0.00030035785,0.029187677,0.00056522747,0.000025290434,0.00023932596,0.00013783263,0.0004721648,0.0046402183],"genre_scores_gemma":[0.98286337,0.00020982609,0.015646249,0.00006469866,0.0000136238905,0.00007503273,0.0002014,0.000060286005,0.0008654865],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97236913,0.010147117,0.0032793472,0.0018637934,0.011442304,0.00089822756],"domain_scores_gemma":[0.74324304,0.13376392,0.060386926,0.02965955,0.03144821,0.0014983558],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018426681,0.00033993102,0.00039384974,0.0042834114,0.0011250435,0.0028696805,0.00129185,0.0012660252,0.0009436678],"category_scores_gemma":[0.15310536,0.00052451435,0.00022301116,0.0029747651,0.0019155717,0.007879465,0.0025301508,0.0015435405,0.00039672156],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049009745,0.00073981023,0.58741385,0.00074371614,0.000058907986,0.0015863302,0.026197683,0.0070784776,0.021197906,0.0073828073,0.0018580265,0.3452524],"study_design_scores_gemma":[0.0000743094,0.00096274505,0.7506134,0.0008195734,0.00014287427,0.004355801,0.021733133,0.08692814,0.08353984,0.017570958,0.032971542,0.00028775935],"about_ca_topic_score_codex":0.0037850942,"about_ca_topic_score_gemma":0.0030598436,"teacher_disagreement_score":0.018426681,"about_ca_system_score_codex":0.0018635038,"about_ca_system_score_gemma":0.003312654,"threshold_uncertainty_score":0.09745079},"labels":[],"label_agreement":null},{"id":"W2027576706","doi":"10.1002/smr.301","title":"Using software trails to reconstruct the evolution of software","year":2004,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software evolution; Computer science; Unix; Software engineering; Documentation; Software development; Software; Source code; Software analytics; Software construction; Software system; Backporting; Process (computing); World Wide Web; Programming language","score_opus":0.06425953317498195,"score_gpt":0.3552294675658492,"score_spread":0.29096993439086727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2027576706","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3202799,0.0015151197,0.6537411,0.00071997737,0.00013312134,0.00023714441,0.0051172418,0.008096475,0.010159871],"genre_scores_gemma":[0.5491372,0.0009095717,0.4365797,0.000048032667,0.000034619807,0.00011997345,0.0072951782,0.00073048205,0.005145272],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99912494,0.00020006225,0.00008241982,0.0001801236,0.00033731363,0.000075154596],"domain_scores_gemma":[0.9916998,0.0019086992,0.0015973352,0.0028896253,0.0016750868,0.00022943622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001633222,0.00042289533,0.00034170251,0.008757657,0.00073565665,0.0022849285,0.0007035107,0.00072665024,0.0027309773],"category_scores_gemma":[0.014065353,0.0005494852,0.0004973977,0.0048094923,0.0009319374,0.0028092633,0.001591122,0.0010217294,0.0016514371],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043741634,0.00024244425,0.10643387,0.00054289325,0.00018410305,0.0013837258,0.006347931,0.07145054,0.017300995,0.039076913,0.0069785872,0.74962056],"study_design_scores_gemma":[0.00009798499,0.0002678685,0.059266917,0.0006017387,0.00018072108,0.0012062041,0.0028475935,0.6934342,0.042297404,0.09184793,0.107741594,0.00020985716],"about_ca_topic_score_codex":0.008980237,"about_ca_topic_score_gemma":0.009705181,"teacher_disagreement_score":0.008980237,"about_ca_system_score_codex":0.0007814871,"about_ca_system_score_gemma":0.0013286384,"threshold_uncertainty_score":0.017855942},"labels":[],"label_agreement":null},{"id":"W2027612103","doi":"10.1109/noms.2012.6212079","title":"Data mining for supporting IT management","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"National Institute for Materials Science; Natural Sciences and Engineering Research Council of Canada; Dalhousie University","keywords":"Computer science; Cluster analysis; Data mining; Focus (optics); Similarity (geometry); Identification (biology); Information retrieval; Knowledge base; Base (topology); Data science; Machine learning; Artificial intelligence","score_opus":0.11793419183923809,"score_gpt":0.3737687556597647,"score_spread":0.2558345638205266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2027612103","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020819204,0.018469198,0.86066127,0.013154791,0.0008644702,0.0022834325,0.0555347,0.009718031,0.018494997],"genre_scores_gemma":[0.11969199,0.008557673,0.814549,0.0015797777,0.00045969788,0.0015469743,0.051065613,0.00023615923,0.0023130323],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9926063,0.0018165545,0.0012085921,0.0013861553,0.002792565,0.00018984487],"domain_scores_gemma":[0.9822463,0.009719009,0.002003718,0.0030420427,0.0025925117,0.00039642924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0078673,0.0015947911,0.0020051836,0.010592455,0.0010833984,0.0045991456,0.0028634,0.001507381,0.0046052462],"category_scores_gemma":[0.025453765,0.0006648201,0.0019185789,0.013983942,0.00069200364,0.004775553,0.0020727806,0.002619128,0.004051799],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035501417,0.00062287616,0.023283016,0.0043610455,0.0008643913,0.00055871706,0.00039334432,0.025142029,0.0058953078,0.05243803,0.055245515,0.83084077],"study_design_scores_gemma":[0.0002373477,0.000387497,0.015671846,0.0027182065,0.0005312406,0.0012418585,0.0012349476,0.284274,0.02315114,0.3015251,0.3687791,0.00024784982],"about_ca_topic_score_codex":0.0033509172,"about_ca_topic_score_gemma":0.003196039,"teacher_disagreement_score":0.010592455,"about_ca_system_score_codex":0.0015748243,"about_ca_system_score_gemma":0.0030549006,"threshold_uncertainty_score":0.041606784},"labels":[],"label_agreement":null},{"id":"W2028793581","doi":"10.1016/j.jss.2011.12.006","title":"The impact of accounting for special methods in the measurement of object-oriented class cohesion on refactoring and fault prediction activities","year":2011,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Kuwait University; University of Alberta","keywords":"Cohesion (chemistry); Code refactoring; Computer science; Empirical research; Class (philosophy); Object-oriented programming; Artificial intelligence; Software; Programming language; Mathematics; Statistics","score_opus":0.06450824714716616,"score_gpt":0.33079920345488045,"score_spread":0.26629095630771427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2028793581","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9856623,0.00032371888,0.012271859,0.00025493623,0.000065452805,0.000017527267,0.00016810177,0.00022485194,0.0010111773],"genre_scores_gemma":[0.9955557,0.000032732834,0.0041553285,0.000022916649,0.000014548689,0.00000492905,0.00008057135,0.0000219378,0.00011136303],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9838754,0.00777692,0.0013612191,0.002124037,0.0039208834,0.00094157836],"domain_scores_gemma":[0.7265895,0.20836781,0.027349012,0.021374017,0.013632428,0.0026872351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01386539,0.00090604153,0.00077205943,0.0018184963,0.00086593145,0.0028892858,0.0011962525,0.0013303707,0.0006518308],"category_scores_gemma":[0.13361593,0.0005366118,0.0006029707,0.0029708194,0.0010364737,0.004130695,0.001181633,0.0016052091,0.0002157062],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023022012,0.00085590885,0.81230664,0.00018138315,0.00052193104,0.00011092541,0.0005642884,0.04470497,0.0072838683,0.0015937175,0.00056884903,0.12900537],"study_design_scores_gemma":[0.000076930606,0.0011682294,0.66385543,0.00008140499,0.00067538075,0.00026438557,0.00042307898,0.31300554,0.016516497,0.0030334105,0.00078763714,0.00011210622],"about_ca_topic_score_codex":0.01603767,"about_ca_topic_score_gemma":0.02211163,"teacher_disagreement_score":0.01603767,"about_ca_system_score_codex":0.0013124638,"about_ca_system_score_gemma":0.0027259185,"threshold_uncertainty_score":0.07332808},"labels":[],"label_agreement":null},{"id":"W2028998897","doi":"10.1109/vissof.2005.1684308","title":"User Perspectives on a Visual Aid to Program Comprehension","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Saint Mary's University; Dalhousie University","funders":"","keywords":"Program comprehension; Computer science; Human–computer interaction; Task (project management); Visualization; Focus (optics); Code (set theory); Source code; Dependency (UML); Comprehension; Distraction; Data visualization; Software engineering; Data science; Software; Programming language; Artificial intelligence; Software system; Set (abstract data type)","score_opus":0.017663068996898776,"score_gpt":0.3288326797805625,"score_spread":0.31116961078366373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2028998897","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.976916,0.0004040393,0.01143982,0.0006212584,0.000040754527,0.00014592739,0.00011139696,0.00047085548,0.009849959],"genre_scores_gemma":[0.9885132,0.00023601086,0.008946671,0.00018152983,0.000030214458,0.0000987629,0.00010076771,0.0001423972,0.0017504883],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99324566,0.005204879,0.00017168568,0.00032698625,0.0007638634,0.00028698987],"domain_scores_gemma":[0.8598303,0.13010925,0.0026324731,0.0024336264,0.003577564,0.0014167087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00898923,0.00096252206,0.0004584831,0.0016230274,0.00087627885,0.003274386,0.00071868085,0.0023106488,0.006158762],"category_scores_gemma":[0.09240233,0.00047774593,0.0005119099,0.00054298213,0.0012767415,0.0031986968,0.001970539,0.0013184161,0.0007651184],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006611195,0.0009932007,0.0401552,0.001618114,0.00009944086,0.0044486364,0.6576133,0.0020176277,0.1589581,0.005612234,0.005431977,0.11644098],"study_design_scores_gemma":[0.0021229787,0.020149061,0.16401826,0.0023967582,0.0009826191,0.013500969,0.467343,0.04491053,0.10392221,0.014782605,0.1644568,0.0014142409],"about_ca_topic_score_codex":0.00091592304,"about_ca_topic_score_gemma":0.00072700245,"teacher_disagreement_score":0.00898923,"about_ca_system_score_codex":0.00051076495,"about_ca_system_score_gemma":0.00036771767,"threshold_uncertainty_score":0.047540188},"labels":[],"label_agreement":null},{"id":"W2029277954","doi":"10.1007/s11219-014-9230-x","title":"Predicting defective modules in different test phases","year":2014,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Benchmark (surveying); Reliability engineering; Computer science; Predictive modelling; Software; Data mining; Machine learning; Engineering; Programming language","score_opus":0.028114770166436563,"score_gpt":0.31264121370854725,"score_spread":0.2845264435421107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2029277954","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9915222,0.00013489599,0.006933991,0.000035607754,0.000012444676,0.000017356364,0.0004883252,0.0005694918,0.00028583308],"genre_scores_gemma":[0.99416757,0.000031839016,0.0041408436,0.000014827473,0.000004826716,0.0000067858255,0.0011646856,0.000047503647,0.00042112538],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987613,0.00020438927,0.00013299522,0.00026915857,0.00042694694,0.0002052227],"domain_scores_gemma":[0.9723652,0.015611279,0.004287153,0.0015386605,0.004799532,0.0013981981],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016448434,0.0009835205,0.00056924776,0.0044500143,0.00025899053,0.00088676723,0.0010891935,0.0012291074,0.0016603012],"category_scores_gemma":[0.01782574,0.0003915818,0.0010376514,0.0014293704,0.00036525086,0.00093338545,0.0004956094,0.00063769065,0.00065198395],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00116196,0.00049726065,0.8908943,0.00009911325,0.000165032,0.00034238107,0.000105703104,0.030315684,0.01134613,0.00020236734,0.0007930036,0.06407714],"study_design_scores_gemma":[0.00006016117,0.0018466464,0.5972244,0.0000493361,0.0003767913,0.001106547,0.00028226717,0.37306926,0.024079992,0.0011519899,0.0006991134,0.000053389427],"about_ca_topic_score_codex":0.005593194,"about_ca_topic_score_gemma":0.0069248034,"teacher_disagreement_score":0.005593194,"about_ca_system_score_codex":0.0005506662,"about_ca_system_score_gemma":0.0005966031,"threshold_uncertainty_score":0.011121273},"labels":[],"label_agreement":null},{"id":"W2029458042","doi":"10.1109/wcre.2013.6671324","title":"An IDE-based context-aware meta search engine","year":2013,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Search engine; Popularity; World Wide Web; Context (archaeology); Eclipse; Information retrieval; Web search engine; Search analytics; Web page; Web search query","score_opus":0.06438170058826184,"score_gpt":0.3158925826069591,"score_spread":0.2515108820186973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2029458042","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12605655,0.0028648763,0.82599735,0.00041086387,0.00027114677,0.00054921804,0.0017611556,0.031204019,0.010884738],"genre_scores_gemma":[0.4294184,0.00062461314,0.55977553,0.00026110667,0.00009766073,0.00024284249,0.0028625673,0.0004423214,0.0062749693],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896526,0.00019363762,0.000118435375,0.00019586261,0.0004322183,0.00009453259],"domain_scores_gemma":[0.9981982,0.0005250388,0.00014404143,0.00037904398,0.0006126512,0.00014099917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011987317,0.00080290023,0.0012605169,0.002761822,0.0005086479,0.0013400721,0.0015895434,0.0010770869,0.0016317928],"category_scores_gemma":[0.003445896,0.00044006627,0.00070039934,0.0015069125,0.00015811174,0.0020469597,0.0011883248,0.00079993147,0.0016082053],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017413687,0.0010661979,0.02154418,0.0009464422,0.00029407983,0.0010000602,0.0005912345,0.011445831,0.07043413,0.0071285116,0.017448299,0.86635965],"study_design_scores_gemma":[0.00043049306,0.00080373563,0.0138273025,0.00017367204,0.00047038583,0.0028074794,0.00058212556,0.86714184,0.06809753,0.0078057083,0.0376266,0.00023306315],"about_ca_topic_score_codex":0.0029374477,"about_ca_topic_score_gemma":0.0076259575,"teacher_disagreement_score":0.0029374477,"about_ca_system_score_codex":0.00030937576,"about_ca_system_score_gemma":0.0011154738,"threshold_uncertainty_score":0.0063396096},"labels":[],"label_agreement":null},{"id":"W2030133704","doi":"10.5555/2819009.2819115","title":"A unified framework for the comprehension of software's time dimension","year":2015,"lang":"en","type":"article","venue":"Papyrus : Institutional Repository (Université de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software evolution; Dimension (graph theory); Program comprehension; Comprehension; Software engineering; Context (archaeology); Software; Software development; Software system; Software construction; Programming language; Mathematics","score_opus":0.014492614360115065,"score_gpt":0.2019449420112597,"score_spread":0.18745232765114464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2030133704","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00046596036,0.00038631182,0.9945932,0.00029044392,0.00006329975,0.0000840782,0.00022300071,0.002046398,0.0018472746],"genre_scores_gemma":[0.016029185,0.0008441609,0.97883314,0.00011297328,0.00010216696,0.00043088372,0.00062498,0.0006598357,0.0023627295],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9964676,0.001208783,0.00042110373,0.0006307115,0.0010137297,0.00025803322],"domain_scores_gemma":[0.9943855,0.0025749174,0.0003768562,0.0011945256,0.0010854302,0.00038277474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059769945,0.002455984,0.0017221466,0.0059330827,0.0018458483,0.010726805,0.004663803,0.002881577,0.0140639255],"category_scores_gemma":[0.01315665,0.0017041529,0.0055855215,0.0040245918,0.0038608338,0.009182871,0.0053721806,0.0045682406,0.0037873315],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011616585,0.00009953321,0.0009783653,0.0010567544,0.00016046029,0.00046793837,0.0032399301,0.03745871,0.0063873613,0.80490154,0.012729476,0.13240388],"study_design_scores_gemma":[0.00006304907,0.00009436369,0.0007493481,0.0007433188,0.00014282504,0.0005094607,0.00081123196,0.25537673,0.0035641077,0.50732,0.23045312,0.00017247173],"about_ca_topic_score_codex":0.017744418,"about_ca_topic_score_gemma":0.018487142,"teacher_disagreement_score":0.017744418,"about_ca_system_score_codex":0.0025546302,"about_ca_system_score_gemma":0.00581186,"threshold_uncertainty_score":0.04704851},"labels":[],"label_agreement":null},{"id":"W2030223557","doi":"10.1016/j.fss.2003.10.007","title":"Building a software experience factory using granular-based models","year":2003,"lang":"en","type":"article","venue":"Fuzzy Sets and Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Institute for Biodiagnostics; University of Alberta","funders":"","keywords":"Computer science; Software development; Software; Visualization; Process (computing); Software engineering; Fuzzy logic; Software development process; Quality (philosophy); Artificial intelligence; Industrial engineering; Data mining; Systems engineering; Engineering","score_opus":0.059816416688458524,"score_gpt":0.2933081701101991,"score_spread":0.23349175342174058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2030223557","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034095712,0.00005062888,0.95997447,0.00011610019,0.000018271317,0.00010373819,0.00008023318,0.0031019675,0.0024588176],"genre_scores_gemma":[0.3864896,0.00013138863,0.610604,0.00004846701,0.000008533945,0.0001758456,0.00025553184,0.00034872664,0.0019379411],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988906,0.0002220868,0.00014350643,0.0002491986,0.00034898863,0.00014571074],"domain_scores_gemma":[0.99819905,0.0005132619,0.00014470836,0.00073760387,0.00022202366,0.00018333057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019529744,0.0006148267,0.001321996,0.00090589526,0.0009138506,0.0046715857,0.0021366654,0.0016566324,0.0044267364],"category_scores_gemma":[0.0043172576,0.0011641912,0.0020991126,0.0009065878,0.0010393452,0.004491718,0.0030185562,0.0017231979,0.0014377458],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006464376,0.00066590676,0.00832779,0.00028523276,0.0002671534,0.0009453569,0.001892535,0.63401395,0.02059592,0.13537332,0.003020109,0.1939663],"study_design_scores_gemma":[0.000045579352,0.00008856212,0.00036567778,0.000032885007,0.000060679024,0.000097679556,0.00014911906,0.9597875,0.0059620813,0.02838614,0.0049854526,0.00003861727],"about_ca_topic_score_codex":0.005298682,"about_ca_topic_score_gemma":0.003744074,"teacher_disagreement_score":0.005298682,"about_ca_system_score_codex":0.0010161705,"about_ca_system_score_gemma":0.0016831735,"threshold_uncertainty_score":0.014808953},"labels":[],"label_agreement":null},{"id":"W2030272387","doi":"10.1016/j.entcs.2004.02.049","title":"DPVK - An Eclipse Plug-in to Detect Design Patterns in Eiffel Systems","year":2004,"lang":"en","type":"article","venue":"Electronic Notes in Theoretical Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Eiffel; Computer science; Eclipse; Software design pattern; Programming language; Software engineering; Reverse engineering; Extensibility; Plug-in; Engineering design process; Compatibility (geochemistry); Architectural pattern; Software; Software design; Software development; Object-oriented programming; Engineering","score_opus":0.01331049559400888,"score_gpt":0.2734982694376952,"score_spread":0.2601877738436863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2030272387","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040848106,0.00015454082,0.7887756,0.00014097964,0.0000892688,0.00032611686,0.0014055873,0.16662867,0.0016311553],"genre_scores_gemma":[0.17346618,0.00022877032,0.8088541,0.00015453133,0.000021632071,0.00039261876,0.0039224043,0.009553288,0.003406458],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984189,0.0003893421,0.00023406488,0.00028644278,0.000538389,0.00013286444],"domain_scores_gemma":[0.9947425,0.0034970448,0.00047140164,0.0007706007,0.00040133906,0.00011714268],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023191771,0.0011723065,0.000645917,0.0017868283,0.00032785206,0.0012789742,0.0014582111,0.0012638145,0.003408826],"category_scores_gemma":[0.010257394,0.0012254204,0.0009389794,0.00064581353,0.00041101582,0.0024271854,0.0013234194,0.0012888267,0.0013450922],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016889597,0.0007992908,0.022294762,0.0019287916,0.00026067658,0.0022259364,0.001474375,0.042716693,0.1295846,0.011205463,0.029636573,0.75618386],"study_design_scores_gemma":[0.0007243032,0.00079091976,0.017347543,0.0005027145,0.00021951293,0.002957618,0.00043967468,0.64278406,0.2342596,0.017471198,0.082267866,0.00023491349],"about_ca_topic_score_codex":0.001355277,"about_ca_topic_score_gemma":0.0018697333,"teacher_disagreement_score":0.003408826,"about_ca_system_score_codex":0.00043451082,"about_ca_system_score_gemma":0.0007308592,"threshold_uncertainty_score":0.012265086},"labels":[],"label_agreement":null},{"id":"W2030573703","doi":"10.1007/s00766-013-0181-8","title":"GaiusT: supporting the extraction of rights and obligations for regulatory compliance","year":2013,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":91,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Health Insurance Portability and Accountability Act; Software portability; Annotation; Compliance (psychology); Computer science; Process (computing); Government (linguistics); Work (physics); Accountability; Computer security; Political science; Law; Engineering; Confidentiality; Artificial intelligence","score_opus":0.04330639118517105,"score_gpt":0.3152901052612124,"score_spread":0.2719837140760414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2030573703","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0084217265,0.00030240795,0.77403855,0.0013492963,0.0004396719,0.0006233653,0.0056476155,0.19651608,0.012661258],"genre_scores_gemma":[0.17728643,0.00047806668,0.7624495,0.0013277376,0.00023235979,0.00062531873,0.019757783,0.02583239,0.012010415],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9898879,0.0035297454,0.00074998103,0.0012323377,0.003856033,0.0007439786],"domain_scores_gemma":[0.9786009,0.011048022,0.0012503992,0.006640731,0.0020386584,0.0004214422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007832424,0.0021730897,0.0013690706,0.0037694082,0.0013093881,0.006538704,0.0036207808,0.0031051536,0.021057075],"category_scores_gemma":[0.045791086,0.001528879,0.0033515366,0.002041552,0.0022137864,0.0074295364,0.008452494,0.0038210189,0.012922059],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015602948,0.0003710518,0.006368659,0.0023982371,0.00053893373,0.0019828503,0.0022309504,0.038490973,0.017465858,0.20852496,0.30139276,0.41867447],"study_design_scores_gemma":[0.0004693215,0.00012764333,0.0011578182,0.0006017539,0.00019796765,0.0008296834,0.00049411954,0.47545788,0.04728249,0.22902265,0.24408251,0.0002760893],"about_ca_topic_score_codex":0.0067984206,"about_ca_topic_score_gemma":0.010657668,"teacher_disagreement_score":0.021057075,"about_ca_system_score_codex":0.001247971,"about_ca_system_score_gemma":0.004133501,"threshold_uncertainty_score":0.070442915},"labels":[],"label_agreement":null},{"id":"W2031271477","doi":"10.1145/2507288.2507326","title":"Technical debt","year":2013,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Pace; Debt; Confusion; Perspective (graphical); Risk analysis (engineering); Software; Computer science; Engineering; Business; Software development; Finance","score_opus":0.014269007103450974,"score_gpt":0.24104295340316437,"score_spread":0.2267739462997134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2031271477","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06211688,0.0047634155,0.04251629,0.026999068,0.0018990259,0.00029564352,0.001125381,0.00060874573,0.8596756],"genre_scores_gemma":[0.735059,0.006821495,0.01158241,0.011831591,0.0014163828,0.0004483764,0.0021013548,0.0006434939,0.23009591],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99012977,0.0023121897,0.0010303656,0.0009815239,0.0043103853,0.0012358192],"domain_scores_gemma":[0.9755803,0.004951181,0.0044313585,0.003861417,0.007997601,0.0031782694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006292099,0.0006082843,0.0004972689,0.0025957946,0.0037343148,0.009670992,0.0017671856,0.0025285978,0.033395894],"category_scores_gemma":[0.037973974,0.00037616963,0.00053884625,0.003700642,0.0040663183,0.010691811,0.0081240935,0.0035563598,0.009665486],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000101228594,0.00010780117,0.012566898,0.00042946765,0.000042383475,0.0010938842,0.008111732,0.0009166542,0.0009872444,0.71978986,0.095556095,0.16029684],"study_design_scores_gemma":[0.000017434479,0.000080046084,0.0063710464,0.0005268823,0.000023475024,0.0021303499,0.003551651,0.0007192389,0.0004408026,0.1725003,0.8135955,0.000043331536],"about_ca_topic_score_codex":0.0025940728,"about_ca_topic_score_gemma":0.0018244113,"teacher_disagreement_score":0.033395894,"about_ca_system_score_codex":0.004810709,"about_ca_system_score_gemma":0.0047920463,"threshold_uncertainty_score":0.11172038},"labels":[],"label_agreement":null},{"id":"W2031927211","doi":"10.1007/s00766-010-0099-3","title":"A controlled experiment to assess the impact of system architectures on new system requirements","year":2010,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Systems engineering; Engineering","score_opus":0.04016249860977132,"score_gpt":0.32888828660906894,"score_spread":0.2887257879992976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2031927211","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98705053,0.00005227577,0.0049317586,0.000114548006,0.00026229952,0.0051783253,0.00043321276,0.00019005273,0.001786967],"genre_scores_gemma":[0.9473751,0.000122961,0.025567453,0.000492095,0.00018584006,0.020132832,0.0007389159,0.00010053986,0.0052841716],"study_design_codex":"nonrandomized_trial","study_design_gemma":"nonrandomized_trial","domain_scores_codex":[0.99507964,0.0017141656,0.0005364952,0.0011468831,0.0009183395,0.0006044249],"domain_scores_gemma":[0.9234727,0.062298406,0.004909779,0.0043500927,0.002567463,0.0024015887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059425873,0.0021093655,0.001271568,0.00066021556,0.0011387717,0.0014266288,0.0023293777,0.0023984564,0.009354043],"category_scores_gemma":[0.023591483,0.0010394873,0.00071366405,0.00042098798,0.0021007005,0.0016559252,0.0012220208,0.0027612625,0.0007735209],"study_design_candidate":"nonrandomized_trial","study_design_consensus":"nonrandomized_trial","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.22828862,0.36325905,0.008527797,0.0024508608,0.00053008687,0.00073996244,0.0044059996,0.01369455,0.3263909,0.0039026504,0.0032757241,0.044533808],"study_design_scores_gemma":[0.113823906,0.7182844,0.023302771,0.00012064819,0.0009385354,0.00017961295,0.0011775936,0.021971358,0.1110643,0.003383379,0.0053617996,0.0003917466],"about_ca_topic_score_codex":0.0012530073,"about_ca_topic_score_gemma":0.0017724608,"teacher_disagreement_score":0.009354043,"about_ca_system_score_codex":0.0011021429,"about_ca_system_score_gemma":0.0027105696,"threshold_uncertainty_score":0.03142774},"labels":[],"label_agreement":null},{"id":"W2032109608","doi":"10.1016/j.ins.2011.07.046","title":"A non-functional requirements tradeoff model in Trustworthy Software","year":2011,"lang":"en","type":"article","venue":"Information Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Natural Science Foundation of China","keywords":"Computer science; Credibility; Fuzzy logic; Data mining; Interdependence; Hierarchy; Software; Sorting; Trustworthiness; Fuzzy set; Operations research; Artificial intelligence; Algorithm; Mathematics; Computer security; Programming language","score_opus":0.09620977499676785,"score_gpt":0.2865736167490604,"score_spread":0.19036384175229254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2032109608","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08918204,0.0003321788,0.8891938,0.0019686471,0.000051916282,0.00012611152,0.000072026676,0.00017742478,0.018895963],"genre_scores_gemma":[0.8942626,0.00017021419,0.10116469,0.00013524688,0.000051574767,0.00014377129,0.00006062538,0.00007148077,0.003939655],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99374455,0.0031262012,0.00026024523,0.0005779015,0.0018864637,0.0004047318],"domain_scores_gemma":[0.97940356,0.015150415,0.0013801077,0.0017713853,0.0018010729,0.0004934632],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008146475,0.0010386227,0.0007766685,0.00183019,0.000894642,0.003435543,0.0025244555,0.0028420745,0.0045927125],"category_scores_gemma":[0.028019842,0.0010606742,0.001082362,0.0010359791,0.0020636958,0.008660873,0.002074781,0.0027036604,0.00064100575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003199093,0.0002114687,0.0022548842,0.00023741602,0.000098957455,0.0004118433,0.0010471495,0.29778007,0.003741378,0.64832014,0.0011814787,0.04439535],"study_design_scores_gemma":[0.000043915545,0.0001322743,0.0007377655,0.00004221512,0.00004660282,0.00015261777,0.00015361013,0.70267254,0.0005411126,0.29416305,0.0012827829,0.000031522766],"about_ca_topic_score_codex":0.0019086057,"about_ca_topic_score_gemma":0.0018055274,"teacher_disagreement_score":0.008146475,"about_ca_system_score_codex":0.0025481563,"about_ca_system_score_gemma":0.0017708171,"threshold_uncertainty_score":0.04308319},"labels":[],"label_agreement":null},{"id":"W2032219610","doi":"10.1109/ms.2009.16","title":"Mining Task-Based Social Networks to Explore Collaboration in Software Teams","year":2008,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":97,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Task (project management); IBM; Social network (sociolinguistics); Context (archaeology); World Wide Web; Software; Software development; Data science; Social network analysis; Software engineering; Social software engineering; Knowledge management; Software construction; Social media; Engineering; Systems engineering","score_opus":0.03378415132012308,"score_gpt":0.2903802627805691,"score_spread":0.25659611146044603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2032219610","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8105038,0.0010744071,0.17488052,0.00092015194,0.00005916553,0.0004333362,0.003586101,0.0004639925,0.008078541],"genre_scores_gemma":[0.9215243,0.00033392024,0.07356216,0.000073169074,0.00005395108,0.0003039112,0.0029634,0.000044561824,0.0011406462],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99773717,0.0011719449,0.00015841215,0.0004168584,0.00037380913,0.00014171156],"domain_scores_gemma":[0.9882533,0.008174675,0.0016123469,0.00088862405,0.0006019125,0.00046924877],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.0019710774,0.0006460454,0.00046698403,0.010979184,0.0010423646,0.0016129898,0.0008440382,0.0009195292,0.0012643425],"category_scores_gemma":[0.012412779,0.00029715305,0.0009430755,0.0063002137,0.0005277779,0.0031361724,0.0017129375,0.00062271167,0.00045204797],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009218222,0.0016468823,0.49024823,0.001197616,0.0013254435,0.0014009323,0.012167043,0.06663002,0.01056302,0.03770465,0.011076408,0.36511797],"study_design_scores_gemma":[0.00009626754,0.00041245233,0.17924868,0.00015859351,0.0003278765,0.0010026047,0.010817763,0.68171865,0.0055285813,0.09825004,0.022306755,0.00013170441],"about_ca_topic_score_codex":0.0033531312,"about_ca_topic_score_gemma":0.006220759,"teacher_disagreement_score":0.99895763,"about_ca_system_score_codex":0.00078992057,"about_ca_system_score_gemma":0.0005723679,"threshold_uncertainty_score":0.010424197},"labels":[],"label_agreement":null},{"id":"W2032227440","doi":"10.1145/2695664.2695938","title":"The safety of dynamic mixin composition","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Modular design; Invariant (physics); Construct (python library); Programming language; Theoretical computer science; Implementation; Upper and lower bounds; Object-oriented programming; Base (topology); Algorithm; Mathematics","score_opus":0.01557055922594152,"score_gpt":0.2701188638805923,"score_spread":0.2545483046546508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2032227440","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07666123,0.00017436166,0.90961826,0.0005491378,0.00008234283,0.00010147465,0.00006575444,0.0017760728,0.010971398],"genre_scores_gemma":[0.8130558,0.00033898844,0.17195047,0.00041793185,0.00016772008,0.00041590413,0.00021179413,0.0009784141,0.012463068],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99306244,0.0013223162,0.00043236365,0.0012590674,0.0031376989,0.0007860865],"domain_scores_gemma":[0.9754431,0.012043057,0.002029779,0.0060310517,0.0036905087,0.0007624276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007197314,0.000768839,0.0010310399,0.0010138709,0.002577673,0.003018796,0.0016820504,0.0016995595,0.0031452614],"category_scores_gemma":[0.024811529,0.0011989635,0.0017905794,0.00045662813,0.006365718,0.006032667,0.005922422,0.0033801894,0.0011199585],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077811565,0.00017927612,0.0040592076,0.00030045974,0.00011026568,0.00080737105,0.0017414708,0.049650144,0.048785802,0.83775723,0.0019891616,0.05384154],"study_design_scores_gemma":[0.00015499583,0.00022756572,0.0006902774,0.000094135146,0.00013476846,0.00046732766,0.00026197374,0.24139707,0.09108384,0.6503263,0.015065668,0.00009609135],"about_ca_topic_score_codex":0.0024335322,"about_ca_topic_score_gemma":0.001003554,"teacher_disagreement_score":0.007197314,"about_ca_system_score_codex":0.0014609596,"about_ca_system_score_gemma":0.0034807438,"threshold_uncertainty_score":0.038063526},"labels":[],"label_agreement":null},{"id":"W2032557572","doi":"10.5555/2820282.2820313","title":"Make it simple: an empirical analysis of GNU make feature use in open source projects","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Computer science; Scripting language; Programming language; Macro; Popularity; Open source; Set (abstract data type); Simple (philosophy); Implementation; Simplicity; Feature (linguistics); Function (biology); Focus (optics); Software engineering; World Wide Web; Software; Linguistics","score_opus":0.1270107041835635,"score_gpt":0.38419889800374174,"score_spread":0.2571881938201782,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2032557572","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9984188,0.00013291258,0.00034395658,0.00009199889,0.000005337856,0.0000150624855,0.00029597207,0.000039084392,0.0006567608],"genre_scores_gemma":[0.99831676,0.00008637378,0.0004601151,0.00003190878,0.000010215461,0.000037091657,0.00067076494,0.00007295319,0.00031365964],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98667496,0.0052770055,0.0012842293,0.001996428,0.004065557,0.0007018951],"domain_scores_gemma":[0.7170713,0.18327428,0.062859155,0.015548707,0.015225583,0.006021079],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011188483,0.0003769139,0.00043229322,0.0064434605,0.0010486037,0.0028314362,0.0012018993,0.0013514496,0.0022270032],"category_scores_gemma":[0.1333248,0.00043326168,0.00054702157,0.006656018,0.0025480264,0.0055153337,0.003072589,0.0017258789,0.0009542229],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001770079,0.000155963,0.97790354,0.000119348784,0.000117649484,0.00016995301,0.0063890303,0.00034035477,0.0005892362,0.00030671322,0.0009766917,0.012754533],"study_design_scores_gemma":[0.0000070664905,0.000118583725,0.99032915,0.00007492602,0.000021399479,0.00040849962,0.005722275,0.0013266318,0.00036044294,0.00021181385,0.0013859653,0.000033224387],"about_ca_topic_score_codex":0.0017405022,"about_ca_topic_score_gemma":0.002205195,"teacher_disagreement_score":0.011188483,"about_ca_system_score_codex":0.0006492504,"about_ca_system_score_gemma":0.00036017896,"threshold_uncertainty_score":0.05917102},"labels":[],"label_agreement":null},{"id":"W2033025900","doi":"10.1109/icsm.2010.5609553","title":"Unit tests as API usage examples","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Documentation; Unit testing; Computer science; Task (project management); Learnability; Software engineering; Presentation (obstetrics); Unit (ring theory); Programming language; Human–computer interaction; Software; Engineering; Systems engineering","score_opus":0.03181308962275948,"score_gpt":0.2961575580967921,"score_spread":0.26434446847403265,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033025900","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42817426,0.0033288118,0.51592624,0.0018736176,0.00040326727,0.0010608588,0.0013185363,0.009432073,0.03848238],"genre_scores_gemma":[0.76932317,0.0010409501,0.21787344,0.0005246132,0.0001541952,0.00073360675,0.0016111945,0.00193612,0.006802837],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9747371,0.01510732,0.0015931488,0.0013096265,0.006565109,0.00068769173],"domain_scores_gemma":[0.80574065,0.15252887,0.012073166,0.0149144,0.013142156,0.0016007712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070345146,0.0013944502,0.0008430588,0.0026551497,0.00057657517,0.0028431003,0.0025348847,0.0019295865,0.0070782513],"category_scores_gemma":[0.12338514,0.0007177678,0.0005704677,0.0021059536,0.0015748431,0.0066795754,0.0025395828,0.0018084826,0.002486596],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00202168,0.001749237,0.040427767,0.005443912,0.0001753516,0.0053848517,0.014499461,0.018744947,0.0346686,0.05981577,0.028379524,0.7886889],"study_design_scores_gemma":[0.0007120072,0.0053031035,0.079733826,0.0068876524,0.0006541256,0.024750043,0.011113206,0.23429292,0.13157082,0.16155554,0.3424245,0.0010022682],"about_ca_topic_score_codex":0.0004907194,"about_ca_topic_score_gemma":0.00060052326,"teacher_disagreement_score":0.0070782513,"about_ca_system_score_codex":0.0005935865,"about_ca_system_score_gemma":0.0006616338,"threshold_uncertainty_score":0.037202477},"labels":[],"label_agreement":null},{"id":"W2033030511","doi":"10.1145/2462932.2462966","title":"ReqWiki","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software engineering; Requirements engineering; Process (computing); Semantic Web; World Wide Web; Requirements analysis; Software; Programming language","score_opus":0.02427330512902074,"score_gpt":0.27669400100480207,"score_spread":0.2524206958757813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033030511","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027348078,0.001090789,0.25665498,0.0020313521,0.0009090601,0.0012549231,0.055856053,0.6013707,0.07809724],"genre_scores_gemma":[0.01988126,0.002638046,0.2925442,0.0030464372,0.00030541632,0.003419912,0.2866493,0.17572959,0.21578583],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9969072,0.0004927958,0.00041599633,0.00072528847,0.0011884847,0.0002702792],"domain_scores_gemma":[0.99357706,0.0017532523,0.0004286527,0.0026470644,0.0011432909,0.00045067153],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032447928,0.00250008,0.0016632685,0.0041600172,0.0018345994,0.0054421206,0.00817793,0.0027369433,0.08278092],"category_scores_gemma":[0.012721396,0.0021906297,0.002509036,0.0032675844,0.0009899097,0.013478639,0.008660263,0.0056920513,0.13687445],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069436984,0.00029124256,0.0005229924,0.0024649592,0.00017018213,0.00046015834,0.0006957453,0.0011058919,0.010632925,0.03811595,0.77120954,0.17363606],"study_design_scores_gemma":[0.0001353388,0.000026560077,0.00033101317,0.00012434309,0.000034022858,0.00025357713,0.000069064605,0.0031151033,0.004834798,0.012381438,0.9785983,0.000096575095],"about_ca_topic_score_codex":0.005292136,"about_ca_topic_score_gemma":0.0076249903,"teacher_disagreement_score":0.08278092,"about_ca_system_score_codex":0.0011609243,"about_ca_system_score_gemma":0.002894241,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2033235985","doi":"10.1109/icsm.2013.34","title":"LHDiff: A Language-Independent Hybrid Approach for Tracking Source Code Lines","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Source code; Source lines of code; Heuristics; Code (set theory); Tracking (education); Programming language; Software; Line (geometry); Artificial intelligence; Computer engineering; Data mining; Operating system","score_opus":0.023990520942369156,"score_gpt":0.2715748712885843,"score_spread":0.24758435034621518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033235985","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021297527,0.0005913711,0.9257194,0.00025037501,0.00012008727,0.0002208347,0.00053630344,0.050158437,0.0011057564],"genre_scores_gemma":[0.09453582,0.00018523104,0.89624965,0.00031392454,0.000050164308,0.00023097688,0.0012215041,0.003634123,0.003578663],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99621165,0.0006184559,0.00039976696,0.00089562207,0.0017044129,0.00017020243],"domain_scores_gemma":[0.9844595,0.0054837144,0.0019453121,0.0050171665,0.0027514782,0.0003428757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002548018,0.0012785479,0.0011201693,0.0039337813,0.0007514348,0.0016711282,0.004210983,0.0015963734,0.0018973589],"category_scores_gemma":[0.010544985,0.0011214531,0.0016045498,0.002434,0.0009789743,0.0034122649,0.0028080668,0.0023818586,0.0014087377],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005222058,0.00031757206,0.010066312,0.0008756186,0.00027630056,0.00043730505,0.001111643,0.016980264,0.10192431,0.0048064804,0.012432673,0.8502493],"study_design_scores_gemma":[0.00031714773,0.0009054369,0.0077796783,0.00017509316,0.00027561287,0.0018348442,0.0005279249,0.6579397,0.24304113,0.0127418535,0.07401379,0.00044781336],"about_ca_topic_score_codex":0.0031721068,"about_ca_topic_score_gemma":0.006030955,"teacher_disagreement_score":0.004210983,"about_ca_system_score_codex":0.00097396126,"about_ca_system_score_gemma":0.0021618681,"threshold_uncertainty_score":0.013475418},"labels":[],"label_agreement":null},{"id":"W2033239109","doi":"10.1145/1101908.1101941","title":"Visualization-based analysis of quality for large-scale software systems","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":169,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Computer science; Visualization; Software visualization; Software quality; Data science; Software evolution; Software; Software development; Software system; Quality (philosophy); Software analytics; Software metric; Software engineering; Visual analytics; Data visualization; Creative visualization; Scale (ratio); Data mining; Software construction","score_opus":0.027724418948046354,"score_gpt":0.3403517475959456,"score_spread":0.31262732864789927,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033239109","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014597446,0.00025410476,0.98229676,0.00029453402,0.00001366401,0.000050171086,0.00007099292,0.0013956876,0.0010265763],"genre_scores_gemma":[0.30535507,0.00028775306,0.6932433,0.00004323658,0.00004865598,0.00009953897,0.00019154938,0.00026256332,0.00046832554],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986199,0.0004492167,0.00009964925,0.00017234028,0.000575946,0.00008299734],"domain_scores_gemma":[0.992184,0.003235649,0.0015828266,0.0012865607,0.0013968276,0.00031414683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025099518,0.00087752007,0.0006872444,0.0053741047,0.00056424405,0.0032355678,0.0011710486,0.0008258664,0.001807104],"category_scores_gemma":[0.010843103,0.00046537316,0.0010916434,0.0023592173,0.0011733186,0.003471758,0.0015100468,0.0012179636,0.00028811407],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019392798,0.00024468053,0.023650534,0.0008529764,0.00033979022,0.0008094771,0.005269134,0.24701717,0.06867287,0.1941995,0.007087431,0.45166242],"study_design_scores_gemma":[0.000033126766,0.00013273029,0.008898387,0.00009867751,0.000071834555,0.00043622992,0.00039432646,0.8749243,0.011068445,0.095874384,0.007973355,0.00009425982],"about_ca_topic_score_codex":0.0022153112,"about_ca_topic_score_gemma":0.0016644123,"teacher_disagreement_score":0.0053741047,"about_ca_system_score_codex":0.0006726592,"about_ca_system_score_gemma":0.0006734507,"threshold_uncertainty_score":0.013274074},"labels":[],"label_agreement":null},{"id":"W2033242799","doi":"10.5555/2664398.2664416","title":"Towards qualitative comparison of simulink model clone detection approaches","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Adaptability; clone (Java method); Relevance (law); Computer science; Visualization; Plan (archaeology); Machine learning; Software engineering; Data mining; Artificial intelligence","score_opus":0.2096295790966588,"score_gpt":0.39586194101826205,"score_spread":0.18623236192160325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033242799","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08307729,0.00019548631,0.910163,0.00012991222,0.000038494934,0.0002180641,0.00021205717,0.0027194633,0.003246369],"genre_scores_gemma":[0.6087425,0.00018006346,0.38894436,0.000041751,0.000010724609,0.00038538664,0.00039512073,0.0003700421,0.0009301016],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9930462,0.0034811457,0.0004537415,0.00044262255,0.0023927048,0.0001835169],"domain_scores_gemma":[0.9545674,0.029105565,0.0030119258,0.004499016,0.00851637,0.00029982542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009609669,0.00080773525,0.00048387947,0.0032218106,0.00042490565,0.0021662856,0.001566589,0.0007267884,0.0040322538],"category_scores_gemma":[0.043274578,0.00036723368,0.0005297061,0.0012213255,0.0010289355,0.0025735325,0.0016518255,0.0006601797,0.0004574742],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026731614,0.00096762273,0.022246623,0.0024960637,0.000300919,0.00040114706,0.004244725,0.24703218,0.12338835,0.08700047,0.0030149797,0.5062338],"study_design_scores_gemma":[0.00020005794,0.0009697475,0.005966678,0.00029898196,0.00012220249,0.00022429155,0.00153807,0.7989954,0.15345447,0.027920807,0.010205085,0.000104181585],"about_ca_topic_score_codex":0.0014783286,"about_ca_topic_score_gemma":0.0010655315,"teacher_disagreement_score":0.009609669,"about_ca_system_score_codex":0.0016026122,"about_ca_system_score_gemma":0.0010971371,"threshold_uncertainty_score":0.050821424},"labels":[],"label_agreement":null},{"id":"W2033771212","doi":"10.1109/msr.2013.6624058","title":"Using citation influence to predict software defects","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Dependency (UML); Computer science; Eclipse; Software; Citation; Component (thermodynamics); Social network analysis; Process (computing); Software development; Data mining; Data science; Software engineering; World Wide Web; Programming language","score_opus":0.02780771782138117,"score_gpt":0.28016165120476816,"score_spread":0.25235393338338696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033771212","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9201986,0.002708214,0.0634943,0.0006307737,0.00016618143,0.00011691811,0.0016338074,0.0016288452,0.009422319],"genre_scores_gemma":[0.989311,0.0005090707,0.008013509,0.000022643939,0.00016013687,0.000040634313,0.0010480243,0.000049202703,0.0008458776],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9982717,0.00040576683,0.00012483643,0.00025670722,0.00080157886,0.00013939613],"domain_scores_gemma":[0.96614397,0.023095395,0.004029846,0.0013190598,0.004521528,0.0008903034],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0029164,0.00090443593,0.00077382463,0.016842695,0.00083185476,0.0014994394,0.00082412944,0.0013917199,0.0012887464],"category_scores_gemma":[0.032750625,0.00034516977,0.0009737854,0.0072009843,0.00045431138,0.0028668833,0.00097504054,0.00080910185,0.0006829883],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004919223,0.00039143895,0.65791786,0.00034164253,0.0007658522,0.00049595005,0.0004967439,0.10215009,0.0043839626,0.0053104083,0.006565348,0.22068876],"study_design_scores_gemma":[0.000032407443,0.00014995899,0.11401014,0.00004965704,0.00034378702,0.00039732773,0.00013175592,0.87028414,0.003639248,0.00807821,0.002826238,0.000057083274],"about_ca_topic_score_codex":0.0140248425,"about_ca_topic_score_gemma":0.015010383,"teacher_disagreement_score":0.9970836,"about_ca_system_score_codex":0.0011338076,"about_ca_system_score_gemma":0.0008244959,"threshold_uncertainty_score":0.02788645},"labels":[],"label_agreement":null},{"id":"W2033890725","doi":"10.1109/esem.2013.18","title":"An Empirical Study of Client-Side JavaScript Bugs","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Unobtrusive JavaScript; JavaScript; Computer science; Rich Internet application; Document Object Model; Web application; Programmer; Server-side; Context (archaeology); Client-side; Empirical research; World Wide Web; Programming language; Web page","score_opus":0.032780765914696045,"score_gpt":0.3267430868070898,"score_spread":0.29396232089239377,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033890725","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9984536,0.00018256853,0.00047363705,0.00010771112,0.0000034283319,0.00010105784,0.0001711687,0.000010402187,0.000496503],"genre_scores_gemma":[0.99799997,0.00023750475,0.00095223583,0.00009659984,0.0000076223846,0.00016994009,0.00031307334,0.000010596226,0.00021244316],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.987225,0.005190599,0.0017033316,0.0011218908,0.0039572367,0.00080196804],"domain_scores_gemma":[0.68517387,0.20326532,0.06698899,0.0067675835,0.03435601,0.003448151],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011386465,0.0004549862,0.0004018333,0.0037050045,0.0009410707,0.0013808294,0.0011379062,0.0010987703,0.0012893989],"category_scores_gemma":[0.10499338,0.00044773048,0.00034328888,0.0036934232,0.0017010488,0.002352126,0.001224294,0.0012559755,0.00036796872],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014468627,0.0010430454,0.9671148,0.00048060482,0.000049117483,0.0005108781,0.012328272,0.00020991004,0.0006415778,0.00023419369,0.000718581,0.01652429],"study_design_scores_gemma":[0.000035751145,0.0013522111,0.96392155,0.00042451025,0.00007667833,0.0012543899,0.026525581,0.0021235272,0.0014760533,0.00024909445,0.0025258388,0.000034912217],"about_ca_topic_score_codex":0.0032446587,"about_ca_topic_score_gemma":0.00443261,"teacher_disagreement_score":0.011386465,"about_ca_system_score_codex":0.0013849211,"about_ca_system_score_gemma":0.0017381718,"threshold_uncertainty_score":0.060218096},"labels":[],"label_agreement":null},{"id":"W2034209539","doi":"10.1109/saner.2015.7081848","title":"CloCom: Mining existing source code for automatic comment generation","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":168,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; KPI-driven code analysis; Code review; Program comprehension; Static program analysis; Source code; Codebase; Programming language; Code (set theory); Redundant code; Source lines of code; Software engineering; Software maintenance; Java; Maintainability; Code generation; Software; Dead code; Software development; Unreachable code; Software system; Operating system; Set (abstract data type)","score_opus":0.15874588872338588,"score_gpt":0.3454253794659993,"score_spread":0.18667949074261345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034209539","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07460733,0.0008598873,0.7130386,0.0009299806,0.0003181665,0.0027665782,0.023072476,0.18029445,0.004112499],"genre_scores_gemma":[0.1237219,0.00037270423,0.80477613,0.00035499982,0.00018056107,0.001875963,0.055009767,0.0081333425,0.0055745607],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99411726,0.0014363196,0.0004611935,0.0013894481,0.0022961868,0.00029956535],"domain_scores_gemma":[0.959591,0.018745417,0.0039876685,0.0050305533,0.011907995,0.0007373281],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048979768,0.0028153828,0.001105118,0.01261723,0.0012198212,0.0015310089,0.0024452135,0.0018759266,0.0042903423],"category_scores_gemma":[0.029225305,0.00086056604,0.0015017478,0.0048999414,0.00084808486,0.003190291,0.002638853,0.0015664734,0.005209022],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008225034,0.00070435955,0.02588609,0.0040237186,0.00029731315,0.0018288507,0.0044067255,0.0056747403,0.09657161,0.0037699873,0.11174514,0.74426895],"study_design_scores_gemma":[0.00053446885,0.0010747524,0.042129304,0.000853477,0.00042351152,0.0029009823,0.0038507637,0.580829,0.18321091,0.011580641,0.17210901,0.00050314394],"about_ca_topic_score_codex":0.004724271,"about_ca_topic_score_gemma":0.0075378995,"teacher_disagreement_score":0.01261723,"about_ca_system_score_codex":0.000945474,"about_ca_system_score_gemma":0.0032188452,"threshold_uncertainty_score":0.025903285},"labels":[],"label_agreement":null},{"id":"W2034259517","doi":"10.1145/1540438.1540453","title":"Practical considerations in deploying AI for defect prediction","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Türkiye Bilimsel ve Teknolojik Araştırma Kurumu; National Aeronautics and Space Administration","keywords":"Computer science; Software bug; Software quality; Code (set theory); Software; Quality (philosophy); Software engineering; Software development; Programming language","score_opus":0.05020378970394145,"score_gpt":0.3554216176059663,"score_spread":0.30521782790202484,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034259517","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2765527,0.0049565285,0.6117314,0.051245507,0.000838396,0.0021158748,0.0005953833,0.012380665,0.039583594],"genre_scores_gemma":[0.5437853,0.0010042131,0.44927555,0.0011620647,0.00027067325,0.00056257506,0.00036620448,0.00040237996,0.0031710588],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9781139,0.016103579,0.0009330365,0.0014391848,0.0028039496,0.0006062873],"domain_scores_gemma":[0.8216308,0.13982351,0.0025113008,0.011859542,0.021613352,0.002561429],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024892466,0.0015377699,0.001088946,0.0019170392,0.000980322,0.004701774,0.0040474413,0.0027534424,0.0054745786],"category_scores_gemma":[0.08380586,0.0010084122,0.00052155723,0.0019877034,0.0014675034,0.0068698004,0.0018542035,0.0028757176,0.002936282],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020344625,0.0017037734,0.06614448,0.0016243028,0.00029795562,0.0018104076,0.003354198,0.10673891,0.031215025,0.0148059055,0.019921409,0.7503492],"study_design_scores_gemma":[0.00066775514,0.002952349,0.031837713,0.00097961,0.00026454302,0.0018198508,0.011201681,0.8352176,0.019911679,0.043771897,0.051069323,0.00030608493],"about_ca_topic_score_codex":0.010246049,"about_ca_topic_score_gemma":0.008120298,"teacher_disagreement_score":0.024892466,"about_ca_system_score_codex":0.0009362306,"about_ca_system_score_gemma":0.002346174,"threshold_uncertainty_score":0.1316455},"labels":[],"label_agreement":null},{"id":"W2034284484","doi":"10.3166/jds.19.201-223","title":"Using Classification Trees to Predict Performance in Information Technology Projects","year":2010,"lang":"en","type":"article","venue":"Journal of Decision System","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Decision tree; Artificial neural network; Machine learning; Binary classification; Artificial intelligence; Data mining; Set (abstract data type); Tree (set theory); Regression; Decision tree learning; Support vector machine; Statistics; Mathematics","score_opus":0.03776090793581022,"score_gpt":0.30533635143018634,"score_spread":0.2675754434943761,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034284484","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76398647,0.0011596057,0.22330515,0.0009153155,0.00013409196,0.00033398933,0.002453357,0.0011192451,0.006592874],"genre_scores_gemma":[0.9386893,0.00044294255,0.05764267,0.000039236507,0.000057213376,0.00023245522,0.0020470638,0.000049205722,0.0007999528],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99594265,0.0021519898,0.00032377255,0.00023920553,0.0010743896,0.00026804316],"domain_scores_gemma":[0.96072614,0.031551834,0.0032956905,0.00084965647,0.0031051456,0.0004715178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057610036,0.0009480932,0.0007126243,0.0058600195,0.0005330073,0.00165568,0.0006250249,0.0010832889,0.0012263766],"category_scores_gemma":[0.03542848,0.00024110195,0.0007063107,0.005902465,0.00037750206,0.0020735422,0.000646384,0.0011982408,0.0006458668],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044570517,0.00044811922,0.27696684,0.00025234607,0.00039667406,0.0001915441,0.00074349006,0.379348,0.0012273474,0.0059918356,0.006411156,0.32757697],"study_design_scores_gemma":[0.00003055517,0.00024823667,0.056266293,0.00010461349,0.00006829165,0.00008614763,0.000319383,0.92539006,0.0012557714,0.014259302,0.0019107141,0.000060652907],"about_ca_topic_score_codex":0.005701678,"about_ca_topic_score_gemma":0.0047085555,"teacher_disagreement_score":0.0058600195,"about_ca_system_score_codex":0.0008607752,"about_ca_system_score_gemma":0.0007726498,"threshold_uncertainty_score":0.03046745},"labels":[],"label_agreement":null},{"id":"W2034357196","doi":"10.1145/2810146.2810152","title":"An Empirical Study of Crash-inducing Commits in Mozilla Firefox","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Crash; Computer science; Programming language","score_opus":0.08375900593016664,"score_gpt":0.37732249724488376,"score_spread":0.29356349131471715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034357196","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99877256,0.00010658637,0.00042342802,0.00011088505,0.000005864728,0.000039521223,0.00015694414,0.000022103886,0.00036218012],"genre_scores_gemma":[0.99817073,0.00012313448,0.0008193757,0.00005105362,0.000007917726,0.000045456858,0.0005029708,0.000016227785,0.00026304965],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9929636,0.0028189046,0.000761708,0.00086937455,0.0021812348,0.00040515087],"domain_scores_gemma":[0.6703568,0.21872082,0.06514325,0.010362049,0.030452749,0.004964251],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011943371,0.00040218263,0.0003481214,0.0028074193,0.000835657,0.0014642571,0.0012966539,0.0009340113,0.0011161668],"category_scores_gemma":[0.149936,0.00050056103,0.00037942544,0.002343004,0.0010114807,0.002456579,0.00087199366,0.0022323795,0.00044263812],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020923007,0.00076651224,0.9765514,0.00014188669,0.00006195857,0.00029730616,0.0057624136,0.00080701325,0.00038905806,0.00021930241,0.0010543413,0.013739565],"study_design_scores_gemma":[0.000028642691,0.0006786213,0.9782606,0.00013831344,0.00004447781,0.0003596251,0.006376283,0.011553985,0.00059109734,0.00029053335,0.001633779,0.00004390407],"about_ca_topic_score_codex":0.01469505,"about_ca_topic_score_gemma":0.017281614,"teacher_disagreement_score":0.01469505,"about_ca_system_score_codex":0.0013929365,"about_ca_system_score_gemma":0.0012006967,"threshold_uncertainty_score":0.06316334},"labels":[],"label_agreement":null},{"id":"W2034400602","doi":"10.1109/quatic.2010.61","title":"IDS: An Immune-Inspired Approach for the Detection of Software Design Smells","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal","funders":"","keywords":"Code smell; Computer science; Creatures; Software engineering; Software; Artificial intelligence; Human–computer interaction; Software quality; Programming language; Software development","score_opus":0.031211199657236635,"score_gpt":0.2650303641087586,"score_spread":0.23381916445152195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034400602","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05648632,0.00077385944,0.93033075,0.00061011116,0.00019374043,0.00026923447,0.00039832722,0.007643344,0.0032943268],"genre_scores_gemma":[0.2967792,0.00028200768,0.69804275,0.00049569184,0.00007989783,0.00018264985,0.0005732988,0.00017665506,0.0033878381],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982052,0.00039025588,0.00013043206,0.0003875937,0.000742895,0.00014365203],"domain_scores_gemma":[0.9977679,0.00077784783,0.00039545455,0.00033195355,0.0005871212,0.00013967593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015678327,0.0010616467,0.0012548221,0.004241581,0.00050787587,0.0013729157,0.0020120544,0.0017108831,0.001092865],"category_scores_gemma":[0.0033994624,0.00046571708,0.0012024237,0.0014885737,0.00079347915,0.0012879518,0.001337987,0.00094222487,0.0006985142],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005192158,0.0010593175,0.028326597,0.00092393986,0.00044563232,0.0009024258,0.00054079824,0.05846063,0.108930774,0.012045196,0.014365434,0.77348006],"study_design_scores_gemma":[0.000045654564,0.00026509396,0.0044364897,0.00004655688,0.000101666505,0.00073015713,0.00011344477,0.9454824,0.030815128,0.009303725,0.008591121,0.00006856138],"about_ca_topic_score_codex":0.0015721276,"about_ca_topic_score_gemma":0.0020341247,"teacher_disagreement_score":0.004241581,"about_ca_system_score_codex":0.0007101836,"about_ca_system_score_gemma":0.00096954,"threshold_uncertainty_score":0.0082915425},"labels":[],"label_agreement":null},{"id":"W2034504975","doi":"10.5555/776816.776922","title":"Panel: empirical validation: what, why, when, and how","year":2003,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; University of Calgary","funders":"","keywords":"Computer science","score_opus":0.0851315501571733,"score_gpt":0.305890984941278,"score_spread":0.22075943478410467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034504975","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062014673,0.02466236,0.042833038,0.68239033,0.00968191,0.0023649407,0.013171615,0.00082909025,0.16205204],"genre_scores_gemma":[0.6917328,0.014069821,0.02670531,0.17969915,0.006762999,0.0048054727,0.026380468,0.0012836127,0.048560373],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.8781745,0.09274881,0.004656558,0.008125767,0.012466623,0.0038278054],"domain_scores_gemma":[0.45914644,0.37909582,0.021195866,0.040430825,0.08972251,0.010408538],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.16729099,0.0009838575,0.0012852277,0.0020793453,0.0036475724,0.0108231595,0.002797227,0.005435307,0.049718086],"category_scores_gemma":[0.4267625,0.00083223206,0.0011471349,0.0049207523,0.004915063,0.011335285,0.004628622,0.005100467,0.013429588],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008518859,0.0003716432,0.09106738,0.0024676034,0.0007846616,0.00021836675,0.003429041,0.0020231092,0.00088958594,0.039569594,0.6649479,0.1933792],"study_design_scores_gemma":[0.0007465467,0.00044305797,0.16290241,0.01433028,0.0009907549,0.0002561097,0.016102992,0.00821478,0.005083562,0.09934529,0.69117755,0.00040671168],"about_ca_topic_score_codex":0.009078974,"about_ca_topic_score_gemma":0.01365146,"teacher_disagreement_score":0.832709,"about_ca_system_score_codex":0.0040735914,"about_ca_system_score_gemma":0.007853708,"threshold_uncertainty_score":0.88472986},"labels":[],"label_agreement":null},{"id":"W2034628356","doi":"10.1007/s10664-005-1290-x","title":"Studying Software Engineers: Data Collection Techniques for Software Field Studies","year":2005,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":491,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; University of Ottawa","funders":"","keywords":"Computer science; Software engineering; Field (mathematics); Software; Task (project management); Data science; Taxonomy (biology); Data collection; Systems engineering; Engineering","score_opus":0.09495362100580118,"score_gpt":0.36248098551184676,"score_spread":0.2675273645060456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034628356","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.141126,0.0014987752,0.6978443,0.0014994371,0.00039385862,0.08746344,0.046235953,0.0030672967,0.020870833],"genre_scores_gemma":[0.13007933,0.00124823,0.6530349,0.0007559146,0.00026201716,0.18638216,0.023060998,0.0006456068,0.004530814],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9390138,0.030056076,0.0122241,0.0044859825,0.012546178,0.0016739687],"domain_scores_gemma":[0.6642607,0.19355331,0.019353226,0.049397107,0.069742106,0.0036935757],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.045168616,0.0020172775,0.0024438668,0.023050098,0.0033682752,0.0031433005,0.003178183,0.0018309392,0.0064847814],"category_scores_gemma":[0.19732668,0.001381583,0.0018618354,0.022674654,0.002129896,0.0042026807,0.0042202366,0.003355481,0.0038778344],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011508809,0.00425472,0.13415684,0.008167575,0.0003675188,0.00046728717,0.02112395,0.003741004,0.01989563,0.017866176,0.062407594,0.7264008],"study_design_scores_gemma":[0.0020502529,0.0043956554,0.42725572,0.0042081946,0.0015523143,0.0010776045,0.045735262,0.030360175,0.08405102,0.05901596,0.33938572,0.0009121765],"about_ca_topic_score_codex":0.0043154215,"about_ca_topic_score_gemma":0.0068305763,"teacher_disagreement_score":0.95483136,"about_ca_system_score_codex":0.0026644887,"about_ca_system_score_gemma":0.0076513425,"threshold_uncertainty_score":0.23887736},"labels":[],"label_agreement":null},{"id":"W2034815413","doi":"10.1109/icsm.2010.5609533","title":"SE-CodeSearch: A scalable Semantic Web-based source code search infrastructure","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; KPI-driven code analysis; Scalability; Source code; Information retrieval; Code (set theory); Context (archaeology); The Internet; Inference; Search engine; Semantic search; Static program analysis; World Wide Web; Software; Programming language; Database; Software development; Artificial intelligence","score_opus":0.01982037521300345,"score_gpt":0.2790464369296371,"score_spread":0.25922606171663365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034815413","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01553708,0.0007755276,0.7083151,0.0010660974,0.00009250087,0.0007786676,0.015822232,0.2456634,0.011949493],"genre_scores_gemma":[0.19055174,0.0009514707,0.7148656,0.00064674334,0.000105902654,0.0008899356,0.07295149,0.013477239,0.0055598677],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972172,0.0004802544,0.0003190893,0.0005563226,0.0012507491,0.000176493],"domain_scores_gemma":[0.99141043,0.0039520343,0.0006361424,0.0021283538,0.0014494808,0.00042360823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031190298,0.0010946763,0.0013119295,0.00873924,0.0011768182,0.0027222058,0.0029539452,0.0016853429,0.009228246],"category_scores_gemma":[0.015018479,0.00079684996,0.0014685743,0.006405985,0.0010307246,0.008648298,0.004313105,0.00177685,0.004954639],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001576749,0.00083278946,0.008957494,0.0026888326,0.0006215252,0.0011885989,0.0014096585,0.022333235,0.03163512,0.07271666,0.23599471,0.62004465],"study_design_scores_gemma":[0.00060721603,0.00028240134,0.005602914,0.00049503596,0.0002462965,0.00078156433,0.00087902555,0.609568,0.042341456,0.11916832,0.21969168,0.00033603248],"about_ca_topic_score_codex":0.0098886285,"about_ca_topic_score_gemma":0.012570871,"teacher_disagreement_score":0.0098886285,"about_ca_system_score_codex":0.0013158702,"about_ca_system_score_gemma":0.0037603208,"threshold_uncertainty_score":0.03087163},"labels":[],"label_agreement":null},{"id":"W2034942038","doi":"10.1002/spip.310","title":"Evaluation of a black‐box estimation tool: A case study","year":2007,"lang":"en","type":"article","venue":"Software Process Improvement and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Benchmarking; Estimation; Computer science; Outlier; White box; Black box; Set (abstract data type); Identification (biology); Software; Data science; Data mining; Industrial engineering; Software engineering; Artificial intelligence; Engineering; Systems engineering","score_opus":0.0475530894494392,"score_gpt":0.3843926142398773,"score_spread":0.3368395247904381,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034942038","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89692074,0.00044724368,0.093138196,0.00067590154,0.0000868924,0.0010386779,0.00061216403,0.0026273227,0.004452798],"genre_scores_gemma":[0.8860027,0.00020181031,0.10891646,0.0001857753,0.000026215374,0.0006094606,0.0010571483,0.0004776483,0.0025228083],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98103034,0.011072001,0.0012127872,0.0016494276,0.004268435,0.0007669735],"domain_scores_gemma":[0.8927947,0.07528137,0.0041653644,0.012889818,0.013072554,0.0017962664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01974974,0.0010095008,0.00077857065,0.002635948,0.0009891787,0.002546609,0.003967292,0.0021915003,0.0025882574],"category_scores_gemma":[0.0611282,0.000601258,0.0007477465,0.0027386204,0.0016331862,0.0035718186,0.0020869898,0.0014099537,0.0011265862],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004887963,0.012293981,0.09616252,0.0034986893,0.0005193389,0.008897671,0.02605118,0.12126222,0.030544376,0.02016396,0.018536536,0.65718156],"study_design_scores_gemma":[0.0011880193,0.013868653,0.074109636,0.0019017737,0.000556912,0.0029368254,0.019433014,0.6993951,0.0888667,0.01088262,0.08634622,0.0005144963],"about_ca_topic_score_codex":0.0033428513,"about_ca_topic_score_gemma":0.003355403,"teacher_disagreement_score":0.01974974,"about_ca_system_score_codex":0.0016154237,"about_ca_system_score_gemma":0.0017525745,"threshold_uncertainty_score":0.1044479},"labels":[],"label_agreement":null},{"id":"W2035011441","doi":"10.1145/373975.373984","title":"Software cost estimation with fuzzy models","year":2000,"lang":"en","type":"article","venue":"ACM SIGAPP Applied Computing Review","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"COCOMO; Fuzzy logic; Fuzzy set; Generalization; Data mining; Computer science; Software; Set (abstract data type); Fuzzy number; Fuzzy classification; Cost estimate; Machine learning; Artificial intelligence; Software development; Mathematics; Engineering; Software construction; Systems engineering; Programming language","score_opus":0.025557746528898178,"score_gpt":0.2739687722956988,"score_spread":0.24841102576680066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2035011441","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023320364,0.00030866437,0.9717459,0.00021661053,0.000021869313,0.000027901371,0.00007328912,0.00009663708,0.00418867],"genre_scores_gemma":[0.87785375,0.0005993894,0.11831682,0.00006906917,0.00007461865,0.00014731006,0.00012996756,0.0000331117,0.00277595],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985411,0.00058622425,0.000056593453,0.00018435293,0.00053208077,0.00009973256],"domain_scores_gemma":[0.99673754,0.0022425845,0.00035578085,0.00019880343,0.0004202296,0.00004507066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019082241,0.00080162665,0.00065067515,0.0017017461,0.00030373316,0.0018891961,0.0011810524,0.0011833592,0.0015036036],"category_scores_gemma":[0.007799034,0.00047649705,0.00089143415,0.0013140155,0.00078902114,0.00167137,0.0007185432,0.00072713377,0.00020490077],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000016093405,0.000010976234,0.00038315027,0.00002404377,0.000022651617,0.000027255166,0.000031380227,0.9619225,0.00020853453,0.025938489,0.00019145111,0.011223452],"study_design_scores_gemma":[0.0000023542339,0.000007968502,0.00011360482,0.000006357632,0.000005140805,0.000007858643,0.000007677747,0.98051065,0.00008888875,0.018973151,0.00027109424,0.0000052389005],"about_ca_topic_score_codex":0.011252229,"about_ca_topic_score_gemma":0.005346034,"teacher_disagreement_score":0.011252229,"about_ca_system_score_codex":0.0020154591,"about_ca_system_score_gemma":0.00082655495,"threshold_uncertainty_score":0.022373438},"labels":[],"label_agreement":null},{"id":"W2035659395","doi":"10.5555/2486788.2486957","title":"Situational awareness: personalizing issue tracking systems","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Situation awareness; Ticket; Context (archaeology); Knowledge management; Tracking system; Human–computer interaction; Software engineering; Process management; World Wide Web; Computer security; Engineering","score_opus":0.03439711888198767,"score_gpt":0.28226766557817023,"score_spread":0.24787054669618258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2035659395","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2937099,0.000933496,0.66331637,0.0025244516,0.00027767764,0.0011000225,0.00029015722,0.02139064,0.016457299],"genre_scores_gemma":[0.76547384,0.00028781442,0.2304481,0.00034127457,0.00014286222,0.0002961987,0.0004100649,0.0005156785,0.002084182],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99202687,0.004633165,0.0005801605,0.0015407589,0.00092505827,0.0002940248],"domain_scores_gemma":[0.9450157,0.03212388,0.005110309,0.011592848,0.0040872353,0.0020699396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011775629,0.0011187429,0.00071482087,0.002776034,0.0013088188,0.0046575787,0.0017048424,0.0014619555,0.0027256473],"category_scores_gemma":[0.05300834,0.0010683114,0.0005298243,0.0013483847,0.0008532026,0.007869481,0.004348474,0.0022182127,0.0009488731],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010175274,0.0016668043,0.06446936,0.0009081664,0.00025833296,0.00056982745,0.025220674,0.014397311,0.03277534,0.009952127,0.010594975,0.8381696],"study_design_scores_gemma":[0.0008158007,0.003618413,0.10255634,0.0014089966,0.0015577404,0.0032542571,0.022305794,0.4872094,0.07478749,0.0685411,0.23287606,0.001068707],"about_ca_topic_score_codex":0.00145273,"about_ca_topic_score_gemma":0.0016209433,"teacher_disagreement_score":0.011775629,"about_ca_system_score_codex":0.00083743903,"about_ca_system_score_gemma":0.0011303519,"threshold_uncertainty_score":0.062276244},"labels":[],"label_agreement":null},{"id":"W2035962249","doi":"10.1109/icsm.2012.6405303","title":"Modelling the &amp;#x2018;Hurried&amp;#x2019; bug report reading process to summarize bug reports","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Process (computing); Point (geometry); Sentence; Reading (process); Software bug; Quality (philosophy); Natural language processing; Information retrieval; Software; Programming language; Linguistics","score_opus":0.05726676498180077,"score_gpt":0.32128686856263633,"score_spread":0.26402010358083555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2035962249","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2077055,0.001025259,0.77990586,0.000967273,0.00011768647,0.0005385197,0.0018574288,0.0058289818,0.0020534701],"genre_scores_gemma":[0.61942655,0.00046230783,0.37003657,0.00012536587,0.00011664207,0.00050553895,0.0037932964,0.0004136977,0.0051200585],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975358,0.0010757257,0.00021817541,0.0006781137,0.00035930655,0.00013274579],"domain_scores_gemma":[0.9803207,0.012080686,0.0033708164,0.0012350727,0.0025340237,0.00045874607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004971679,0.0011790325,0.0007840736,0.0023772956,0.00047079546,0.0027146447,0.0011889979,0.0014765061,0.0019574566],"category_scores_gemma":[0.02849011,0.0006058695,0.0009006007,0.0018156082,0.0005434479,0.0023512968,0.0009699168,0.0013478756,0.0011079547],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017507091,0.0008747112,0.07056366,0.0013099926,0.0004968054,0.000589357,0.0063625886,0.421099,0.030370789,0.011809564,0.013153367,0.44161943],"study_design_scores_gemma":[0.000038888058,0.00017438427,0.007294099,0.000029739887,0.00007663726,0.00007404754,0.00014877945,0.9833686,0.0034141091,0.0030915968,0.0022496006,0.000039584815],"about_ca_topic_score_codex":0.013539407,"about_ca_topic_score_gemma":0.017199649,"teacher_disagreement_score":0.013539407,"about_ca_system_score_codex":0.0011147403,"about_ca_system_score_gemma":0.0013168512,"threshold_uncertainty_score":0.026921213},"labels":[],"label_agreement":null},{"id":"W2036164091","doi":"10.1109/re.2007.55","title":"Viewing Project Collaborators WhoWork on Interrelated Requirements","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Requirements analysis; Computer science; Requirements engineering; Requirements management; Requirement prioritization; Set (abstract data type); Plan (archaeology); Non-functional requirement; Listing (finance); Interdependence; Work breakdown structure; Work (physics); Project management; Software requirements specification; Visualization; Software project management; Project team; Project management triangle; Software; Software development; Systems engineering; Knowledge management; Engineering; Project charter; Software design","score_opus":0.045076553923334255,"score_gpt":0.3373198162288613,"score_spread":0.29224326230552705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2036164091","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3393839,0.0025327844,0.38386032,0.005923222,0.00085832825,0.0009443547,0.007760867,0.017531885,0.24120437],"genre_scores_gemma":[0.72766626,0.0017259684,0.21974982,0.00044698402,0.00019656261,0.00049174554,0.0047174897,0.00317996,0.041825213],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99930453,0.0002915863,0.000037678747,0.00011053381,0.00019583723,0.000059844548],"domain_scores_gemma":[0.9962202,0.0019114704,0.00030265536,0.00036568532,0.00076094735,0.00043907398],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014231252,0.0007962578,0.00039290605,0.0030165652,0.0009567641,0.0018766759,0.00071043975,0.0010535351,0.02342384],"category_scores_gemma":[0.007215425,0.0003616033,0.00033150453,0.0023426283,0.0003053929,0.0032184217,0.0024250941,0.0009736518,0.0033785626],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001435036,0.00038547372,0.02952259,0.0014947872,0.00010823045,0.004509936,0.16705804,0.005474333,0.050640088,0.076552644,0.17299755,0.48982123],"study_design_scores_gemma":[0.00026231282,0.0004125284,0.034684975,0.001090136,0.00018621069,0.0026179713,0.048951592,0.02444621,0.01654716,0.025597477,0.8449288,0.00027464965],"about_ca_topic_score_codex":0.0040263063,"about_ca_topic_score_gemma":0.00532117,"teacher_disagreement_score":0.02342384,"about_ca_system_score_codex":0.00051127933,"about_ca_system_score_gemma":0.0009664632,"threshold_uncertainty_score":0.07836056},"labels":[],"label_agreement":null},{"id":"W2036333174","doi":"10.1145/351936.351953","title":"Choosing an object-oriented domain framework","year":2000,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Citation; Domain (mathematical analysis); Object (grammar); Library science; Artificial intelligence; Mathematics","score_opus":0.052600460995895056,"score_gpt":0.3584658570058323,"score_spread":0.30586539600993723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2036333174","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017466256,0.38716823,0.33630788,0.052582327,0.002219278,0.00088835176,0.00019826362,0.0011569083,0.20201246],"genre_scores_gemma":[0.15096033,0.3214617,0.45082554,0.016750095,0.001447469,0.0008065707,0.00067473884,0.00042863292,0.05664497],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9944352,0.0016908224,0.00031021095,0.00048717236,0.002656911,0.00041963372],"domain_scores_gemma":[0.9965939,0.0011843429,0.00036001156,0.00029343803,0.0012072471,0.00036109262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008753085,0.0005730314,0.00064822,0.0030399936,0.0017356508,0.005023621,0.0021188648,0.0026984417,0.0033917169],"category_scores_gemma":[0.007915159,0.00064692023,0.0006002251,0.0028938856,0.0017114931,0.007933759,0.0020033761,0.0022902011,0.004034361],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004949376,0.00012556235,0.0015728627,0.0016762976,0.000036002428,0.000489975,0.000708644,0.0011658176,0.002437035,0.29124874,0.034476805,0.66601276],"study_design_scores_gemma":[0.000037927417,0.00007567769,0.0009851923,0.0025322952,0.00004112826,0.0010814484,0.0012432631,0.0015766999,0.0012518059,0.11650165,0.8746136,0.000059324186],"about_ca_topic_score_codex":0.004583396,"about_ca_topic_score_gemma":0.004852057,"teacher_disagreement_score":0.008753085,"about_ca_system_score_codex":0.0029702247,"about_ca_system_score_gemma":0.0050014895,"threshold_uncertainty_score":0.04629135},"labels":[],"label_agreement":null},{"id":"W2036389624","doi":"10.1109/esem.2013.7","title":"Message from the PROMISE 2013 Chairs","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software quality assurance; Quality assurance; Software engineering; Process (computing); Verifiable secret sharing; Engineering management; Software quality; Quality (philosophy); Software; Software project management; Software development; Data science; Software construction; Engineering","score_opus":0.013524611762063243,"score_gpt":0.22772803490741406,"score_spread":0.21420342314535082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2036389624","genre_codex":"commentary","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00048362094,0.002030515,0.0006132169,0.8506334,0.13196006,0.000078479105,0.00082931726,0.0003325476,0.01303892],"genre_scores_gemma":[0.008157838,0.0018827699,0.0011926404,0.76396465,0.0736443,0.0004336363,0.00075009954,0.0005113531,0.14946267],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99221134,0.0013368384,0.00029037322,0.0010859399,0.003985684,0.0010897558],"domain_scores_gemma":[0.98246896,0.0023937437,0.00056021713,0.0006592559,0.0073997444,0.0065180627],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014575885,0.0014987619,0.001630014,0.0009869857,0.0061451904,0.009191561,0.0032677008,0.020397522,0.062490165],"category_scores_gemma":[0.02711168,0.0007741595,0.002460945,0.000959277,0.0020934031,0.0060350266,0.0053668334,0.038450923,0.04906672],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022632332,0.000013097104,0.000043020773,0.000014774295,0.0000018173948,0.000032879674,0.000018528643,0.000009460039,0.000035015986,0.0006428546,0.9978136,0.0013521997],"study_design_scores_gemma":[0.000022244767,0.000029514165,0.00056014716,0.000080833386,0.000007796184,0.000068209214,0.00019301272,0.00009621189,0.00013730906,0.0007311904,0.99803835,0.000035184774],"about_ca_topic_score_codex":0.0112502985,"about_ca_topic_score_gemma":0.013289192,"teacher_disagreement_score":0.062490165,"about_ca_system_score_codex":0.005488059,"about_ca_system_score_gemma":0.011600751,"threshold_uncertainty_score":0.20905042},"labels":[],"label_agreement":null},{"id":"W2036487649","doi":"10.1109/ms.2012.24","title":"Contemporary Peer Review in Action: Lessons from Open Source Development","year":2012,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":134,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; University of Victoria; Concordia University","funders":"Engineering and Physical Sciences Research Council","keywords":"Agile software development; Code review; Software engineering; Computer science; Software development; Software quality; Software inspection; Software development process; Software peer review; Personal software process; Process (computing); Technical peer review; Team software process; Asynchronous communication; Software technical review; Software; Process management; Engineering; Software construction; Peer review; Operating system","score_opus":0.16966384960637565,"score_gpt":0.3791838742844237,"score_spread":0.20952002467804803,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2036487649","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013530224,0.13357046,0.059310943,0.6559158,0.008357276,0.0002131217,0.000050423678,0.0006234655,0.12842825],"genre_scores_gemma":[0.5969913,0.1787229,0.06473467,0.08828709,0.028619835,0.00084173674,0.00011654415,0.0012841215,0.040401716],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.8849905,0.063046396,0.004861034,0.006559724,0.037598733,0.0029435786],"domain_scores_gemma":[0.5870537,0.287298,0.014760427,0.023755206,0.071827136,0.015305552],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.096789055,0.00070824416,0.0011770334,0.004738653,0.0075835623,0.020582868,0.0046179397,0.010883195,0.004731936],"category_scores_gemma":[0.19360618,0.00087476225,0.00082904665,0.0042520487,0.028540555,0.022591654,0.010385295,0.01029564,0.0027428379],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011876686,0.00024445725,0.0041533425,0.0030510337,0.000101788006,0.0011403484,0.028075948,0.0016743654,0.0005655588,0.2762411,0.12514001,0.55949336],"study_design_scores_gemma":[0.00009841439,0.00017947335,0.0037871082,0.0022088035,0.000031530577,0.0010485055,0.01174201,0.0011620368,0.00052945653,0.2344897,0.7445889,0.00013401256],"about_ca_topic_score_codex":0.005146816,"about_ca_topic_score_gemma":0.005521915,"teacher_disagreement_score":0.90321094,"about_ca_system_score_codex":0.00786183,"about_ca_system_score_gemma":0.015853953,"threshold_uncertainty_score":0.51187557},"labels":[],"label_agreement":null},{"id":"W2036570657","doi":"10.1007/s00165-009-0115-x","title":"From a domain analysis to the specification and detection of code and design smells","year":2009,"lang":"en","type":"article","venue":"Formal Aspects of Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal","funders":"","keywords":"Code smell; Computer science; Context (archaeology); Domain (mathematical analysis); Adaptation (eye); Domain analysis; Software engineering; Programming language; Software system; Software; Software development; Software quality; Software construction","score_opus":0.015197072992355601,"score_gpt":0.24887866242311502,"score_spread":0.23368158943075942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2036570657","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15900984,0.00035136586,0.8298387,0.0002946749,0.00001927697,0.0001393224,0.000195954,0.008914248,0.0012364953],"genre_scores_gemma":[0.50646627,0.00025894554,0.4906851,0.00015234764,0.000015990252,0.0001198167,0.00073978945,0.0007221814,0.0008395178],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99392384,0.0028873621,0.0005213122,0.00084225356,0.0015862617,0.0002388958],"domain_scores_gemma":[0.96467906,0.02153185,0.0036455013,0.0064846445,0.0032972137,0.00036174373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040538334,0.00094351673,0.0008888854,0.0030970562,0.0004285831,0.0019346615,0.000987153,0.0011518218,0.00087847974],"category_scores_gemma":[0.021162236,0.000746831,0.0012494695,0.0010115654,0.0009688322,0.0016734448,0.0017023821,0.0013858665,0.00055667036],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006787065,0.0007070034,0.04320674,0.0013585002,0.0003148684,0.0015026103,0.002290097,0.10765183,0.16359642,0.016290277,0.004517309,0.65788555],"study_design_scores_gemma":[0.00009296991,0.00032681462,0.009137727,0.00020873672,0.00013572359,0.0010406944,0.00051284354,0.78773683,0.17905994,0.013136397,0.008511772,0.00009958038],"about_ca_topic_score_codex":0.0009787827,"about_ca_topic_score_gemma":0.0009023357,"teacher_disagreement_score":0.0040538334,"about_ca_system_score_codex":0.00076837454,"about_ca_system_score_gemma":0.0013162118,"threshold_uncertainty_score":0.021438956},"labels":[],"label_agreement":null},{"id":"W2036661337","doi":"10.1145/1159733.1159766","title":"A family of empirical studies to compare informal and optimization-based planning of software releases","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software release life cycle; Replication (statistics); Computer science; Empirical research; Software; Process (computing); Operations research; Process management; Management science; Software quality; Software development; Engineering; Mathematics; Statistics","score_opus":0.04905786676383975,"score_gpt":0.33702278248609063,"score_spread":0.2879649157222509,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2036661337","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7223356,0.0035540643,0.24925223,0.0010207101,0.00051103346,0.007372144,0.0013391088,0.0007680789,0.013847005],"genre_scores_gemma":[0.89789826,0.001195502,0.09115797,0.0006422719,0.00022047767,0.00670207,0.0008440782,0.00016893618,0.0011704864],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.93190014,0.045211434,0.0067148875,0.006281324,0.008760705,0.0011314532],"domain_scores_gemma":[0.25789768,0.62499803,0.03350353,0.061906762,0.020127898,0.0015660917],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.047648776,0.0013436562,0.0012985546,0.0026525322,0.0013953275,0.0016029631,0.002751699,0.0019985675,0.0030884482],"category_scores_gemma":[0.36887473,0.0008429893,0.0016105912,0.0033445454,0.0027337486,0.004641758,0.0019879767,0.003027711,0.000503272],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.014199139,0.035596695,0.13672432,0.015021165,0.005851297,0.00081964553,0.019972999,0.050457425,0.023525188,0.031708676,0.007922469,0.65820104],"study_design_scores_gemma":[0.008931097,0.13623346,0.44663307,0.006523396,0.005360971,0.0032980363,0.02430928,0.13154313,0.08457751,0.075943284,0.07489022,0.0017565379],"about_ca_topic_score_codex":0.0023447336,"about_ca_topic_score_gemma":0.0019729412,"teacher_disagreement_score":0.9523512,"about_ca_system_score_codex":0.0021975373,"about_ca_system_score_gemma":0.0018790102,"threshold_uncertainty_score":0.25199383},"labels":[],"label_agreement":null},{"id":"W2036780516","doi":"10.1007/s00236-009-0107-6","title":"Early action in an Earley parser","year":2009,"lang":"en","type":"article","venue":"Acta Informatica","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Parsing; Compiler; Natural language processing; Artificial intelligence; Programming language; Action (physics); Theory of computation; LR parser; Top-down parsing","score_opus":0.027492417611835438,"score_gpt":0.2901093124590566,"score_spread":0.2626168948472211,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2036780516","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05127667,0.00068186794,0.8460914,0.002538695,0.0007371977,0.00019116428,0.002124063,0.067181915,0.029176971],"genre_scores_gemma":[0.45520943,0.00069300196,0.47787666,0.0014095941,0.0005457345,0.00014461455,0.0029116946,0.026692774,0.034516487],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99531806,0.0012355224,0.00046256397,0.0010933597,0.0013306058,0.00055993174],"domain_scores_gemma":[0.9841287,0.011414105,0.00039850784,0.0021478508,0.0016859657,0.00022481277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067261653,0.0017617816,0.0021041126,0.0020829777,0.0018198184,0.005298122,0.0035684078,0.0037506127,0.02514048],"category_scores_gemma":[0.014692061,0.0032334987,0.0014556181,0.0026649425,0.002303035,0.011147766,0.004176717,0.0042699026,0.007914979],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001984278,0.00025699276,0.004191349,0.0010334189,0.00013212173,0.0031029484,0.004551397,0.010465204,0.03406305,0.49617627,0.06395022,0.3800928],"study_design_scores_gemma":[0.0002847592,0.00022638825,0.0017961946,0.00038749693,0.0005580411,0.0016133632,0.00094605505,0.16095303,0.17154051,0.5073641,0.15387776,0.0004522679],"about_ca_topic_score_codex":0.0063484376,"about_ca_topic_score_gemma":0.006265264,"teacher_disagreement_score":0.02514048,"about_ca_system_score_codex":0.0019370302,"about_ca_system_score_gemma":0.003475239,"threshold_uncertainty_score":0.084103346},"labels":[],"label_agreement":null},{"id":"W2037019767","doi":"10.1016/j.ins.2004.08.012","title":"Early estimation of software size in object-oriented environments a case study in a CMM level 3 software firm","year":2005,"lang":"en","type":"article","venue":"Information Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"European Commission; University of Calgary","keywords":"Computer science; Software; Software sizing; Scope (computer science); Software metric; Software measurement; Software development; Verification and validation; Software engineering; Software quality; Software construction; Statistics; Programming language; Mathematics","score_opus":0.030821818046954802,"score_gpt":0.29896613642638853,"score_spread":0.26814431837943375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2037019767","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99317557,0.000085367006,0.004839061,0.00009365303,0.0000025451875,0.000026739468,0.00005690922,0.000028442037,0.0016917522],"genre_scores_gemma":[0.99474996,0.000039882518,0.0044792164,0.000009253973,0.0000030335764,0.000009630559,0.00007388786,0.000010510264,0.00062459585],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99785006,0.00090333296,0.00014994363,0.00020130187,0.0007041917,0.00019117513],"domain_scores_gemma":[0.92074805,0.06307381,0.004754476,0.0024274176,0.007937048,0.0010591712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033932857,0.00031125257,0.00029003998,0.0022550044,0.00066034583,0.0010954577,0.0012844781,0.0009780141,0.0013460213],"category_scores_gemma":[0.028218241,0.00030688755,0.0003067372,0.001769512,0.0005408528,0.0017893618,0.00076262304,0.0008185796,0.00021468641],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009040535,0.0014228896,0.79432994,0.00024682126,0.00004913355,0.0039408333,0.012150499,0.013044455,0.015144065,0.00480788,0.0008489189,0.15311043],"study_design_scores_gemma":[0.000040533283,0.0017704598,0.8339455,0.00014058939,0.00011431264,0.0015922973,0.01365934,0.11589071,0.024404105,0.0039645964,0.0043513784,0.00012620512],"about_ca_topic_score_codex":0.014219511,"about_ca_topic_score_gemma":0.024020614,"teacher_disagreement_score":0.014219511,"about_ca_system_score_codex":0.0017059075,"about_ca_system_score_gemma":0.0007888768,"threshold_uncertainty_score":0.028273523},"labels":[],"label_agreement":null},{"id":"W2037556552","doi":"10.5555/776816.776826","title":"Using benchmarking to advance research: a challenge to software engineering","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":200,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Toronto","funders":"","keywords":"Benchmarking; Benchmark (surveying); Computer science; Data science; Software; Software engineering; Academic community; Management science; Engineering management; Engineering","score_opus":0.112132774218325,"score_gpt":0.3623551205264038,"score_spread":0.25022234630807877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2037556552","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020345636,0.026467884,0.6693411,0.24609712,0.0037965274,0.0005994729,0.0003300469,0.0034498516,0.02957244],"genre_scores_gemma":[0.21866761,0.014605792,0.74339944,0.012624796,0.0030804318,0.0019237282,0.00073785544,0.002158564,0.0028017454],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.6318311,0.24535716,0.02013659,0.014170409,0.08441986,0.0040849047],"domain_scores_gemma":[0.30232066,0.45884833,0.026822725,0.111394174,0.08751724,0.013097004],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.27593943,0.0029740357,0.0050456487,0.011453028,0.0055520097,0.027796308,0.010212177,0.008602041,0.0017468493],"category_scores_gemma":[0.5178069,0.0012971504,0.001566891,0.017814683,0.019904891,0.0545129,0.018991953,0.015850464,0.0021934845],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011595696,0.00031814913,0.005439621,0.0017574824,0.00016631211,0.00016552806,0.0037541026,0.01088346,0.0011816667,0.4884381,0.033942774,0.45383677],"study_design_scores_gemma":[0.00007579211,0.0003725815,0.0014168868,0.001985407,0.00004561293,0.00015858772,0.0028673562,0.015852822,0.0015029848,0.8726366,0.10289488,0.00019049174],"about_ca_topic_score_codex":0.0027222899,"about_ca_topic_score_gemma":0.0024577253,"teacher_disagreement_score":0.72406054,"about_ca_system_score_codex":0.006429577,"about_ca_system_score_gemma":0.022274131,"threshold_uncertainty_score":0.8928956},"labels":[],"label_agreement":null},{"id":"W2037602363","doi":"10.1007/s00766-014-0214-y","title":"Requirements for tools for comprehending highly specialized assembly language code and how to elicit these requirements","year":2014,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software engineering; Program comprehension; Malware; Reverse engineering; Requirements analysis; Domain (mathematical analysis); Visualization; Cyberspace; Software; Government (linguistics); Data science; Computer security; World Wide Web; Software system; The Internet; Programming language; Artificial intelligence","score_opus":0.06456112806267894,"score_gpt":0.327813087020144,"score_spread":0.26325195895746506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2037602363","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010692082,0.00008983844,0.9755351,0.0010437734,0.00003875136,0.0012605671,0.0006271946,0.0034909786,0.007221734],"genre_scores_gemma":[0.075740226,0.00017622985,0.91516197,0.00033297465,0.000030717292,0.001561022,0.0025725882,0.0018053565,0.002618854],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.97647095,0.007624471,0.0034456463,0.0012317355,0.010173739,0.0010533641],"domain_scores_gemma":[0.8658966,0.08809643,0.0070085996,0.014462229,0.023217859,0.0013182365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014081331,0.0023918664,0.0013088569,0.0034502724,0.0015878751,0.005427214,0.003036686,0.004498668,0.0067978282],"category_scores_gemma":[0.114379786,0.0024220634,0.0021598584,0.0018726572,0.0023712837,0.0073300716,0.0038719226,0.0047296034,0.0043111513],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053576985,0.0007926355,0.0058034747,0.00577373,0.0001385001,0.005699226,0.01668172,0.052565902,0.17573094,0.33136228,0.03309996,0.3718158],"study_design_scores_gemma":[0.00029240685,0.0007763768,0.006815103,0.00449634,0.00034297898,0.007932151,0.0094645,0.32595888,0.18867505,0.20165946,0.25312984,0.00045685348],"about_ca_topic_score_codex":0.002685135,"about_ca_topic_score_gemma":0.003116482,"teacher_disagreement_score":0.014081331,"about_ca_system_score_codex":0.0014068466,"about_ca_system_score_gemma":0.0050571063,"threshold_uncertainty_score":0.0744701},"labels":[],"label_agreement":null},{"id":"W2037905258","doi":"10.1145/1454474.1454485","title":"Dynamic analysis of Ada programs for comprehension and quality measurement","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Program comprehension; Call graph; Software quality; Schema (genetic algorithms); Software engineering; Visualization; Programming language; Comprehension; Software; Software system; Data mining; Software development; Information retrieval","score_opus":0.11941379264931433,"score_gpt":0.32900320970754593,"score_spread":0.2095894170582316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2037905258","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5212212,0.00051159156,0.45346767,0.00027242035,0.000035613615,0.0002591461,0.0014034673,0.019264488,0.0035645233],"genre_scores_gemma":[0.8337778,0.00021608178,0.16221571,0.00004818186,0.000025715623,0.00021742829,0.001510832,0.0011215094,0.0008667641],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975817,0.0005113871,0.00018400868,0.00038991263,0.0012208057,0.00011219641],"domain_scores_gemma":[0.98916686,0.004743847,0.0014030156,0.0016725428,0.0027901342,0.0002235361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018225375,0.00076086196,0.0005894543,0.003580745,0.00035253135,0.0013150057,0.00060350687,0.0004049813,0.0010199065],"category_scores_gemma":[0.010446964,0.00029284833,0.00048062784,0.0020758042,0.0004851563,0.0014711777,0.000628141,0.00086909946,0.00033625244],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007984584,0.00072252407,0.0944535,0.0009211843,0.00019236124,0.0005565849,0.00443068,0.038828727,0.3180602,0.00986991,0.0037103097,0.52745557],"study_design_scores_gemma":[0.00007039127,0.0013000348,0.12873136,0.00016326473,0.00022330417,0.0010081804,0.001131022,0.56874543,0.26820138,0.011474263,0.018769216,0.00018214365],"about_ca_topic_score_codex":0.0015648492,"about_ca_topic_score_gemma":0.0013764474,"teacher_disagreement_score":0.003580745,"about_ca_system_score_codex":0.0005866991,"about_ca_system_score_gemma":0.00067191664,"threshold_uncertainty_score":0.0096386075},"labels":[],"label_agreement":null},{"id":"W2038172274","doi":"10.1109/icsm.2013.79","title":"gCad: A Near-Miss Clone Genealogy Extractor to Support Clone Evolution Analysis","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Extractor; Computer science; Programming language; Biology; Genetics; Engineering; Gene","score_opus":0.015450807150031845,"score_gpt":0.27077041797515977,"score_spread":0.25531961082512794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2038172274","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04820361,0.00080112036,0.56622404,0.00031600066,0.000100692756,0.00038391963,0.014678543,0.36690494,0.0023870643],"genre_scores_gemma":[0.18099253,0.00032206043,0.77380514,0.00033754026,0.000087060325,0.0006474567,0.025178079,0.013913004,0.004717098],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981059,0.00015443873,0.00023603957,0.0005241368,0.000866272,0.000113186834],"domain_scores_gemma":[0.9905059,0.0039685895,0.0014999391,0.0019191806,0.0017796893,0.00032673893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019150814,0.0014284695,0.001059641,0.00902984,0.00076213735,0.001989633,0.0022024873,0.0014263884,0.0047859163],"category_scores_gemma":[0.015109767,0.00081109925,0.0010180343,0.0049716365,0.00066142855,0.0032385078,0.002410434,0.0012563476,0.0025774732],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008857146,0.0002680749,0.052792367,0.00095407816,0.00032796452,0.00069481065,0.0017089441,0.010585896,0.031875573,0.006121087,0.07413517,0.8196504],"study_design_scores_gemma":[0.00029072512,0.0004262454,0.041375503,0.0002727549,0.00025866012,0.0017136163,0.0007234219,0.6580319,0.14034538,0.021294465,0.13494264,0.0003248343],"about_ca_topic_score_codex":0.005400058,"about_ca_topic_score_gemma":0.008677204,"teacher_disagreement_score":0.00902984,"about_ca_system_score_codex":0.0010533655,"about_ca_system_score_gemma":0.0015522564,"threshold_uncertainty_score":0.016010523},"labels":[],"label_agreement":null},{"id":"W2038637109","doi":"10.5555/2664360.2664388","title":"Systematic mapping of recommendation systems for requirements engineering","year":2012,"lang":"en","type":"article","venue":"International Conference on Software and System Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Requirements engineering; Process (computing); Requirements analysis; Recommender system; Requirements management; User requirements document; Functional requirement; Quality (philosophy); Systems engineering; Product (mathematics); Work (physics); Software engineering; Engineering; Information retrieval; Software","score_opus":0.08301229068590288,"score_gpt":0.330319403439364,"score_spread":0.24730711275346112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2038637109","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045927517,0.002397057,0.9202826,0.0007793442,0.00009184149,0.0023952415,0.0018090035,0.004152685,0.022164594],"genre_scores_gemma":[0.1667697,0.0014592718,0.82427686,0.000108546876,0.000016221638,0.001650364,0.0031137706,0.00044192237,0.0021632526],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9822686,0.009281323,0.0019733754,0.0015819956,0.004333267,0.00056144415],"domain_scores_gemma":[0.9572753,0.018977942,0.002602976,0.008325347,0.012448626,0.00036985026],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013915656,0.001409122,0.0011178928,0.015633065,0.0019766272,0.004722332,0.0012308119,0.001321415,0.004552763],"category_scores_gemma":[0.05084985,0.0014189774,0.0024001182,0.012190382,0.0011087497,0.0052991845,0.0027944033,0.0012376362,0.0016805242],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035248036,0.00030646578,0.017021494,0.0039007799,0.00041759873,0.00054524624,0.007504365,0.02121709,0.009368074,0.052237894,0.007073498,0.8800551],"study_design_scores_gemma":[0.00041727244,0.0019959593,0.06618038,0.007260495,0.0013346148,0.0030161375,0.01735036,0.34748825,0.061206073,0.18628258,0.30660543,0.00086243043],"about_ca_topic_score_codex":0.012107609,"about_ca_topic_score_gemma":0.0102643045,"teacher_disagreement_score":0.98608434,"about_ca_system_score_codex":0.0028178,"about_ca_system_score_gemma":0.007097719,"threshold_uncertainty_score":0.073593915},"labels":[],"label_agreement":null},{"id":"W2038876950","doi":"10.1080/08839510802028447","title":"CRAWLING THE CONSTRUCTION WEB – A MACHINE-LEARNING APPROACH WITHOUT NEGATIVE EXAMPLES","year":2008,"lang":"en","type":"article","venue":"Applied Artificial Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Web crawler; Crawling; Web page; Domain (mathematical analysis); Focused crawler; Class (philosophy); World Wide Web; Artificial intelligence; Machine learning; Information retrieval; Web navigation; Static web page","score_opus":0.06203447223025554,"score_gpt":0.2727079318439812,"score_spread":0.21067345961372563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2038876950","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07827849,0.0013645067,0.9107724,0.0012948492,0.00007831231,0.000393101,0.00054474885,0.00328276,0.003990924],"genre_scores_gemma":[0.3514553,0.00044933733,0.63902533,0.00052784354,0.00021540422,0.00042847844,0.0020447744,0.00020649243,0.0056470768],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963058,0.0014118012,0.00031389712,0.00102622,0.00081280083,0.0001293853],"domain_scores_gemma":[0.98602504,0.009151369,0.0008068622,0.00196224,0.001870002,0.00018458845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036479707,0.0012557307,0.0014936136,0.0043205176,0.0012222741,0.0023268047,0.0029265077,0.002844627,0.0010065534],"category_scores_gemma":[0.014514775,0.0008274328,0.0009676763,0.00217071,0.0017983031,0.0032641462,0.001385982,0.0016267418,0.00095101306],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026649833,0.0010803336,0.01510433,0.00096442056,0.00022825657,0.0007442308,0.0012596865,0.07028374,0.018304083,0.0129581075,0.014148978,0.8646574],"study_design_scores_gemma":[0.000059718583,0.0002108113,0.0044202907,0.0001304273,0.000092191265,0.00084539986,0.0002893634,0.9356293,0.012660166,0.03466192,0.010943998,0.000056467616],"about_ca_topic_score_codex":0.0026501995,"about_ca_topic_score_gemma":0.0037904652,"teacher_disagreement_score":0.0043205176,"about_ca_system_score_codex":0.0007405194,"about_ca_system_score_gemma":0.0009119123,"threshold_uncertainty_score":0.019292533},"labels":[],"label_agreement":null},{"id":"W2038917214","doi":"10.1016/j.infsof.2014.12.006","title":"ELBlocker: Predicting blocking bugs with ensemble imbalance learning","year":2015,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":98,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"National Key Research and Development Program of China","keywords":"Blocking (statistics); Computer science; Software bug; Disjoint sets; Eclipse; Artificial intelligence; Software; Programming language; Mathematics; Computer network","score_opus":0.008737206462084155,"score_gpt":0.22172110999079375,"score_spread":0.2129839035287096,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2038917214","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21330976,0.004665253,0.65996057,0.0012673899,0.0015448446,0.0005266573,0.013007742,0.1006344,0.005083488],"genre_scores_gemma":[0.572726,0.0006048687,0.38377854,0.0007653544,0.0005961473,0.00047868807,0.029234616,0.0018943481,0.009921382],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99859613,0.00031457434,0.00009297752,0.00039254947,0.00040808058,0.0001957305],"domain_scores_gemma":[0.99558216,0.0022293597,0.0004000145,0.0007733779,0.000776095,0.00023912397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024365284,0.0023552992,0.0015056769,0.0030378886,0.00061503734,0.0009305018,0.0026064336,0.0018232985,0.0044176504],"category_scores_gemma":[0.0076406375,0.0006175352,0.0010895863,0.001511656,0.00032221864,0.0025673246,0.0016459525,0.0025460883,0.002370361],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014840598,0.0011673647,0.027949011,0.00034412587,0.00063608825,0.00027366308,0.00009545893,0.14731447,0.008514855,0.0017970995,0.098129645,0.71229416],"study_design_scores_gemma":[0.00008262327,0.00014865029,0.001202326,0.000015899182,0.00006957318,0.000048485676,0.00001559308,0.989963,0.0032695844,0.0024085296,0.0027583293,0.00001740984],"about_ca_topic_score_codex":0.004701707,"about_ca_topic_score_gemma":0.0067711174,"teacher_disagreement_score":0.004701707,"about_ca_system_score_codex":0.0005812606,"about_ca_system_score_gemma":0.0012961068,"threshold_uncertainty_score":0.014778554},"labels":[],"label_agreement":null},{"id":"W2039052699","doi":"10.1007/s10664-012-9228-6","title":"Studying re-opened bugs in open source software","year":2012,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":138,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Software bug; Eclipse; Software regression; Dimension (graph theory); Open source; Rework; Computer science; Software; Software quality; Software engineering; Engineering; Software development; Operating system; Mathematics","score_opus":0.05923128895567467,"score_gpt":0.32261791254277666,"score_spread":0.26338662358710196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039052699","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99785393,0.0001623411,0.0011506287,0.00010115244,0.0000041818416,0.000008830975,0.000027357179,0.000018001023,0.0006734681],"genre_scores_gemma":[0.99822384,0.00010575931,0.0009825827,0.000022579601,0.0000065005506,0.000008775873,0.000111088135,0.000025320871,0.000513525],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9965423,0.0013531797,0.00028895805,0.000429693,0.0010822909,0.00030362935],"domain_scores_gemma":[0.8241104,0.12869783,0.024886852,0.008768903,0.011185049,0.0023510705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046089008,0.00035011285,0.00032809438,0.002582762,0.000849801,0.001363404,0.0011074973,0.0011209928,0.0020648893],"category_scores_gemma":[0.10787618,0.00042104317,0.00040849167,0.0021825014,0.0011373146,0.0041203625,0.0013390464,0.0018862034,0.00023852078],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003537371,0.0020977797,0.8862586,0.00030938679,0.00023673216,0.0007326549,0.01151197,0.0037664045,0.004047874,0.005303929,0.0010836805,0.084297225],"study_design_scores_gemma":[0.00006626383,0.0011094655,0.9323799,0.00026409072,0.00019900955,0.0011947152,0.017681899,0.026572244,0.004678101,0.012496293,0.0032880763,0.000069960064],"about_ca_topic_score_codex":0.006470883,"about_ca_topic_score_gemma":0.011809086,"teacher_disagreement_score":0.006470883,"about_ca_system_score_codex":0.0010018197,"about_ca_system_score_gemma":0.0011527039,"threshold_uncertainty_score":0.024374485},"labels":[],"label_agreement":null},{"id":"W2039083805","doi":"10.1145/1321211.1321227","title":"A comparative study of pairwise regression techniques for problem determination","year":2007,"lang":"en","type":"article","venue":"Proceedings of CASCON","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Pairwise comparison; Computer science; Regression; Regression analysis; Artificial intelligence; Statistics; Machine learning; Mathematics","score_opus":0.03592794746935532,"score_gpt":0.3429632541745326,"score_spread":0.3070353067051773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039083805","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0146917915,0.0017249346,0.9779544,0.00038958265,0.000095319585,0.00013989536,0.00010682559,0.0026550198,0.0022423028],"genre_scores_gemma":[0.12630185,0.0010808186,0.86852735,0.00014042652,0.00009610749,0.00022357256,0.00046531824,0.0009945114,0.002170079],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97729003,0.013921414,0.0009462228,0.0022371209,0.005081106,0.0005241352],"domain_scores_gemma":[0.873616,0.10408432,0.0031189371,0.008566259,0.009817865,0.00079662167],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014989466,0.0021761297,0.0020061743,0.0040076543,0.0009286573,0.0012861266,0.0035442242,0.001393152,0.006510352],"category_scores_gemma":[0.1091131,0.0006549285,0.0022602847,0.00493163,0.0008777509,0.004197974,0.0027441527,0.0031978884,0.0020178326],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004588133,0.000382362,0.0028733925,0.0007632059,0.00036012777,0.00007729624,0.0004912281,0.12495851,0.0035096495,0.016766323,0.004919226,0.84443974],"study_design_scores_gemma":[0.000059180282,0.00037847937,0.0012899679,0.000057948968,0.00010648682,0.00020836027,0.00033283434,0.9732702,0.0046838834,0.014753866,0.0048197503,0.00003903327],"about_ca_topic_score_codex":0.0039878893,"about_ca_topic_score_gemma":0.0047020325,"teacher_disagreement_score":0.014989466,"about_ca_system_score_codex":0.0011766287,"about_ca_system_score_gemma":0.0024017538,"threshold_uncertainty_score":0.07927281},"labels":[],"label_agreement":null},{"id":"W2039142846","doi":"10.1145/1454247.1454256","title":"Understanding interaction differences between newcomer and expert programmers","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Human–computer interaction","score_opus":0.26962014628892966,"score_gpt":0.33283886433438215,"score_spread":0.06321871804545248,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039142846","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99231124,0.00009504896,0.0061018686,0.000092634444,0.0000050405483,0.000019878344,0.000015801568,0.000030004103,0.0013285244],"genre_scores_gemma":[0.997083,0.000037359383,0.0023469578,0.000026902591,0.0000082494935,0.000022352899,0.00004350199,0.00001580136,0.00041590547],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9937634,0.0029304093,0.0003552643,0.0010167753,0.001474568,0.00045967603],"domain_scores_gemma":[0.9090457,0.07029,0.007860654,0.004164391,0.005711858,0.0029273534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065091457,0.00020564611,0.00038752434,0.0015310716,0.00071847846,0.0017909475,0.0009286175,0.0007742823,0.0018158808],"category_scores_gemma":[0.07456847,0.00024846688,0.00022224996,0.0005184119,0.00085500686,0.0028908344,0.0023679796,0.0008149058,0.00022201781],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018759479,0.00040508245,0.6358349,0.0003078562,0.00016116384,0.0011555153,0.1822107,0.0010648497,0.046960138,0.003245441,0.0005752934,0.12620303],"study_design_scores_gemma":[0.00010681055,0.0014338192,0.88438433,0.00007931251,0.00017151413,0.001984947,0.077305935,0.011807515,0.010750424,0.007897739,0.003934283,0.00014339798],"about_ca_topic_score_codex":0.0015356555,"about_ca_topic_score_gemma":0.001699909,"teacher_disagreement_score":0.0065091457,"about_ca_system_score_codex":0.0005379933,"about_ca_system_score_gemma":0.0005020059,"threshold_uncertainty_score":0.034424007},"labels":[],"label_agreement":null},{"id":"W2039240270","doi":"10.1109/vissoft.2014.28","title":"ChronoTwigger: A Visual Analytics Tool for Understanding Source and Test Co-evolution","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Visual analytics; Visualization; Pace; Source code; Software analytics; Analytics; Software visualization; Data science; Cultural analytics; Process (computing); Software; Software development; Software engineering; Human–computer interaction; Data mining; Software development process; Information retrieval; Software construction; Semantic analytics; Programming language","score_opus":0.030733418405548885,"score_gpt":0.2927087011250041,"score_spread":0.2619752827194552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039240270","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016014146,0.0006122197,0.8178554,0.00089001865,0.00018444947,0.0005431476,0.015409763,0.14201167,0.0064792335],"genre_scores_gemma":[0.12649815,0.0009065848,0.8420917,0.00030976496,0.00009781297,0.0014163833,0.011906307,0.012133244,0.0046401448],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99896336,0.0002924072,0.00012872452,0.00017319071,0.00036251973,0.000079844904],"domain_scores_gemma":[0.9914601,0.0059129903,0.00061222824,0.00090202456,0.0008064391,0.0003062229],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038281805,0.0015665054,0.00078946794,0.007895002,0.0007557131,0.0034484712,0.001655081,0.0012206383,0.021302795],"category_scores_gemma":[0.012821057,0.0007870739,0.0011561303,0.0041554463,0.00067336764,0.0051326915,0.0034236654,0.0022148958,0.0026233965],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016777251,0.00043059833,0.012859664,0.0034644043,0.00041408502,0.0012589515,0.012063164,0.022869539,0.029235814,0.0464963,0.16597863,0.7032511],"study_design_scores_gemma":[0.00061572954,0.00041850784,0.016093012,0.0013496837,0.00023028108,0.0013953735,0.0036359744,0.40773523,0.0381871,0.11850179,0.41129935,0.0005379608],"about_ca_topic_score_codex":0.004488566,"about_ca_topic_score_gemma":0.005660789,"teacher_disagreement_score":0.021302795,"about_ca_system_score_codex":0.00086172955,"about_ca_system_score_gemma":0.0017092273,"threshold_uncertainty_score":0.07126492},"labels":[],"label_agreement":null},{"id":"W2039316517","doi":"10.1142/s0218194012400116","title":"DEFECT PREDICTION USING CASE-BASED REASONING: AN ATTRIBUTE WEIGHTING TECHNIQUE BASED UPON SENSITIVITY ANALYSIS IN NEURAL NETWORKS","year":2012,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Weighting; Artificial neural network; Sensitivity (control systems); Data mining; Computer science; Heuristic; Artificial intelligence; Linear regression; Machine learning; Pattern recognition (psychology); Engineering","score_opus":0.01592908508520493,"score_gpt":0.26951147551447335,"score_spread":0.25358239042926844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039316517","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042644717,0.00019996442,0.9546546,0.00013536916,0.000030352874,0.00014847529,0.00009319602,0.00065260136,0.0014406942],"genre_scores_gemma":[0.6894485,0.00022864371,0.30901057,0.00008601535,0.000036024307,0.00017561711,0.00022006598,0.000058249927,0.00073634036],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974899,0.00079783733,0.00022149427,0.00040935312,0.0009559883,0.00012536788],"domain_scores_gemma":[0.99335676,0.0043732417,0.0006339352,0.0004971337,0.0010466498,0.000092273054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036247615,0.0010647182,0.0009183948,0.0040843263,0.00043938283,0.0011100507,0.0010860624,0.00088504015,0.0014071186],"category_scores_gemma":[0.015188792,0.00046507298,0.0014477388,0.002094206,0.0005277573,0.0017476586,0.0009335254,0.00092439266,0.00021713247],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016418261,0.00026249498,0.008045782,0.0001744942,0.00030935454,0.00026686906,0.0002544373,0.6706335,0.008399494,0.0048068697,0.000889962,0.30579254],"study_design_scores_gemma":[0.0000052822247,0.000031138672,0.00093518215,0.000016859825,0.00003513739,0.00005510323,0.00001791559,0.9928076,0.0021656791,0.0036796115,0.0002359599,0.0000145266795],"about_ca_topic_score_codex":0.004360502,"about_ca_topic_score_gemma":0.0027535511,"teacher_disagreement_score":0.004360502,"about_ca_system_score_codex":0.0009011471,"about_ca_system_score_gemma":0.00067851663,"threshold_uncertainty_score":0.019169807},"labels":[],"label_agreement":null},{"id":"W2039536418","doi":"","title":"Discovering Information Explaining API Types Using Text Classification","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Application programming interface; Interface (matter); Precision and recall; Information retrieval; Software; Artificial intelligence; Programming language","score_opus":0.03867956524393828,"score_gpt":0.27813749512913816,"score_spread":0.23945792988519987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039536418","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50969756,0.0038768565,0.43502638,0.0012666266,0.0004517291,0.0008972455,0.014578302,0.023247425,0.010957872],"genre_scores_gemma":[0.69117945,0.00082066155,0.2750353,0.0002539412,0.00037673808,0.0004002248,0.026090128,0.00040466883,0.005438896],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987771,0.00028620445,0.00014103211,0.00040453358,0.0002992788,0.00009176696],"domain_scores_gemma":[0.9853496,0.008985782,0.0015842895,0.0009219068,0.0028387646,0.00031962822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014835881,0.0014329794,0.0006243379,0.008601504,0.0006852791,0.0015686324,0.0010087303,0.0014338434,0.0031860962],"category_scores_gemma":[0.010706538,0.000261028,0.0010482091,0.0032392163,0.0003229435,0.002558565,0.0006644939,0.0011369387,0.0033132648],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056085276,0.00087682024,0.051882204,0.0007489038,0.00019330635,0.0005595713,0.00049628725,0.010980168,0.03840671,0.0011546598,0.029677855,0.86446273],"study_design_scores_gemma":[0.000105950545,0.0005278563,0.06670153,0.00047263943,0.0003871465,0.0011927966,0.0009633231,0.8181262,0.07268314,0.008259656,0.030449674,0.0001299503],"about_ca_topic_score_codex":0.0036482222,"about_ca_topic_score_gemma":0.003528867,"teacher_disagreement_score":0.008601504,"about_ca_system_score_codex":0.000650279,"about_ca_system_score_gemma":0.00091024506,"threshold_uncertainty_score":0.010658562},"labels":[],"label_agreement":null},{"id":"W2039729958","doi":"10.1002/spip.298","title":"Using Data Envelopment Analysis in software development productivity measurement","year":2006,"lang":"en","type":"article","venue":"Software Process Improvement and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Benchmarking; Data envelopment analysis; Productivity; Software; Computer science; Function point; Software development; Function (biology); Regression analysis; Process (computing); Linear regression; Industrial engineering; Engineering; Statistics; Business; Mathematics; Economics; Machine learning","score_opus":0.09050916239602833,"score_gpt":0.3356399786467488,"score_spread":0.2451308162507205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039729958","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08353436,0.006840033,0.89688486,0.001138435,0.0001108534,0.00034785946,0.00085090235,0.00024695435,0.01004569],"genre_scores_gemma":[0.69237196,0.0076473616,0.29678968,0.00020413636,0.00007897683,0.0008186354,0.000909821,0.000091234506,0.0010880982],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.95937884,0.027729169,0.0022447205,0.0013781339,0.00843155,0.00083761645],"domain_scores_gemma":[0.9255171,0.0597519,0.005259058,0.0047602244,0.004389789,0.0003219232],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.030748367,0.0018389339,0.0020856098,0.00903638,0.0009201259,0.004910394,0.0011044471,0.0017440808,0.00076834235],"category_scores_gemma":[0.13610415,0.00074243074,0.001569286,0.026965134,0.001812707,0.004381781,0.0038641014,0.0022876018,0.00038126518],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022871746,0.00031632572,0.032063622,0.0011547994,0.00064300135,0.00027670988,0.0015964853,0.58644444,0.001971901,0.13593031,0.0028871323,0.23648655],"study_design_scores_gemma":[0.000047929345,0.0004925999,0.016291792,0.00079823344,0.00017871405,0.0001685672,0.0013707586,0.81033856,0.005236218,0.14931375,0.015500365,0.00026249007],"about_ca_topic_score_codex":0.020493198,"about_ca_topic_score_gemma":0.008495901,"teacher_disagreement_score":0.96925163,"about_ca_system_score_codex":0.0047243317,"about_ca_system_score_gemma":0.0062509347,"threshold_uncertainty_score":0.16261488},"labels":[],"label_agreement":null},{"id":"W2039822794","doi":"10.1145/1767751.1767754","title":"Clone region descriptors","year":2010,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"clone (Java method); Computer science; Cloning (programming); Source code; Software maintenance; Software; Code refactoring; Code (set theory); Software system; Programming language; Biology; Genetics; Gene","score_opus":0.08537287773891615,"score_gpt":0.3123188283960462,"score_spread":0.22694595065713002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039822794","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06564231,0.002512052,0.80405194,0.00059327245,0.00066492934,0.0017319659,0.063005015,0.033005066,0.028793523],"genre_scores_gemma":[0.2774572,0.0019082067,0.5981359,0.00073448924,0.00026968218,0.002337646,0.092723146,0.0052692196,0.02116448],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9975586,0.0002459388,0.00052287686,0.00053157797,0.0009818113,0.0001592272],"domain_scores_gemma":[0.9899761,0.0032101432,0.0013744816,0.002448106,0.002688449,0.00030266235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018182928,0.0008652869,0.0011191554,0.006439158,0.0008188445,0.0028055962,0.0019412681,0.001610651,0.0097695235],"category_scores_gemma":[0.012638841,0.0005281951,0.0010165089,0.006371964,0.00093455426,0.004146131,0.002059548,0.0012154175,0.005075682],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012214324,0.00021448573,0.0327211,0.0017976689,0.00011991272,0.0012932075,0.0027115464,0.010363257,0.041991208,0.15762453,0.09656881,0.6533728],"study_design_scores_gemma":[0.00020451631,0.00044220372,0.014872849,0.0004110702,0.00015445771,0.0031500582,0.00097061065,0.048803426,0.051117826,0.071919695,0.8076757,0.00027751183],"about_ca_topic_score_codex":0.003828177,"about_ca_topic_score_gemma":0.0031539686,"teacher_disagreement_score":0.0097695235,"about_ca_system_score_codex":0.0013544243,"about_ca_system_score_gemma":0.0017454588,"threshold_uncertainty_score":0.03268236},"labels":[],"label_agreement":null},{"id":"W2039854894","doi":"10.1145/2568225.2568270","title":"Effects of using examples on structural model comprehension: a controlled experiment","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Comprehension; Domain (mathematical analysis); Unified Modeling Language; Class diagram; Domain model; Program comprehension; Empirical research; Completeness (order theory); Class (philosophy); Artificial intelligence; Software; Software engineering; Domain knowledge; Programming language; Software system","score_opus":0.02814075456058587,"score_gpt":0.2881083396112961,"score_spread":0.25996758505071027,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039854894","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.987949,0.00010132976,0.005625758,0.00014171323,0.000060883296,0.004475964,0.00025057816,0.00015460922,0.0012400054],"genre_scores_gemma":[0.9302612,0.00017716769,0.044359487,0.0004891271,0.00011136601,0.02164488,0.0004673392,0.00011853821,0.0023709848],"study_design_codex":"bench_or_experimental","study_design_gemma":"nonrandomized_trial","domain_scores_codex":[0.98997974,0.005614136,0.0010423348,0.0017389801,0.0009750046,0.00064981217],"domain_scores_gemma":[0.8172142,0.15939559,0.00806898,0.0074129836,0.0051463787,0.0027618825],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012109799,0.0019456984,0.0013470023,0.00058306253,0.001103034,0.002084715,0.0023744171,0.0028354025,0.008500988],"category_scores_gemma":[0.06469156,0.0012896205,0.00081721647,0.0004813647,0.0016998281,0.0024872168,0.0021808976,0.0024567167,0.0008249132],"study_design_candidate":"nonrandomized_trial","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.10779404,0.3105063,0.021153968,0.007860293,0.0006698149,0.0017389024,0.056709778,0.011444056,0.3388614,0.0029484588,0.004403016,0.13591005],"study_design_scores_gemma":[0.07204023,0.58398414,0.08684579,0.0010444318,0.0017954991,0.0008514248,0.0098531665,0.027359534,0.1867729,0.0067720083,0.021707155,0.00097375974],"about_ca_topic_score_codex":0.000523742,"about_ca_topic_score_gemma":0.00061322784,"teacher_disagreement_score":0.012109799,"about_ca_system_score_codex":0.0007011349,"about_ca_system_score_gemma":0.0013443219,"threshold_uncertainty_score":0.06404352},"labels":[],"label_agreement":null},{"id":"W2039978418","doi":"10.1007/s10515-011-0098-8","title":"Maintainability defects detection and correction: a multi-objective approach","year":2012,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":144,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Université de Montréal","funders":"","keywords":"Code refactoring; Maintainability; Computer science; Sorting; Reliability engineering; Genetic programming; Software; Process (computing); Code (set theory); Software engineering; Programming language; Machine learning; Set (abstract data type); Engineering","score_opus":0.010863922017511712,"score_gpt":0.23886260324286082,"score_spread":0.2279986812253491,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2039978418","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025091816,0.0003524057,0.97245026,0.00014046536,0.000022114398,0.00017339863,0.000087182496,0.00039383103,0.0012884727],"genre_scores_gemma":[0.41964805,0.00025513154,0.57605356,0.00010693345,0.00006215704,0.0003918721,0.0003208239,0.00018203128,0.002979441],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99461716,0.0014935661,0.00043749332,0.0007828306,0.002295629,0.00037337551],"domain_scores_gemma":[0.9923029,0.0033278686,0.0012801508,0.00033684072,0.002512793,0.00023945773],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066624302,0.002977318,0.0032097825,0.0076137846,0.0010444138,0.0027719382,0.0028111357,0.0020591188,0.0017721453],"category_scores_gemma":[0.007749815,0.0011186097,0.0025866146,0.002452761,0.00093606074,0.002287221,0.0023570624,0.0011499894,0.0003132056],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003315667,0.00068275124,0.01039254,0.00056460703,0.0010148506,0.00019203743,0.00032781548,0.57226616,0.016466644,0.004884131,0.001045985,0.39183083],"study_design_scores_gemma":[0.000020890324,0.00026557955,0.0023088965,0.00003242386,0.00015991763,0.00004990082,0.00007688255,0.9911329,0.0029289643,0.0026235243,0.00036618568,0.0000338173],"about_ca_topic_score_codex":0.0062295226,"about_ca_topic_score_gemma":0.0072323494,"teacher_disagreement_score":0.0076137846,"about_ca_system_score_codex":0.0016151038,"about_ca_system_score_gemma":0.0025650507,"threshold_uncertainty_score":0.03523469},"labels":[],"label_agreement":null},{"id":"W2040027946","doi":"10.1002/smr.413","title":"Recommending change clusters to support software investigation: an empirical study","year":2009,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Software; Empirical research; Software system; Software maintenance; Change impact analysis; Software development; Source code; Software engineering; Source lines of code; Code (set theory); Programming language; Set (abstract data type)","score_opus":0.14764349092597215,"score_gpt":0.42174833325183925,"score_spread":0.27410484232586707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2040027946","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9961671,0.00012451362,0.002328085,0.00018086174,0.000005946124,0.00033929176,0.00010038851,0.000105196435,0.0006484563],"genre_scores_gemma":[0.98747665,0.00012732377,0.011620003,0.00006501604,0.000012006761,0.00023179255,0.00020283161,0.000030255285,0.00023414941],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97973734,0.013136452,0.0015072201,0.0016465373,0.0033941253,0.00057835126],"domain_scores_gemma":[0.4403803,0.4738624,0.0384277,0.019576177,0.021665005,0.006088381],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02368205,0.00044946326,0.0006175675,0.005072901,0.0017807971,0.002888264,0.0022113707,0.0016969143,0.0018758221],"category_scores_gemma":[0.22729753,0.00076825154,0.000424008,0.0034826035,0.0012197846,0.0041672415,0.0021607166,0.0020143078,0.00051656907],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015037217,0.007611031,0.6924777,0.0011128798,0.00017696057,0.0009485007,0.04293118,0.0028344842,0.0039202445,0.0010448848,0.0036705849,0.24176784],"study_design_scores_gemma":[0.00054485013,0.005810713,0.8643814,0.0006745268,0.0003675286,0.0015962645,0.047526136,0.05569553,0.006659792,0.0016448799,0.014834419,0.00026403437],"about_ca_topic_score_codex":0.0031112365,"about_ca_topic_score_gemma":0.004235642,"teacher_disagreement_score":0.02368205,"about_ca_system_score_codex":0.0017738999,"about_ca_system_score_gemma":0.0025100806,"threshold_uncertainty_score":0.12524414},"labels":[],"label_agreement":null},{"id":"W2040225071","doi":"10.1002/spip.146","title":"Findings from Phase 2 of the SPICE trials","year":2001,"lang":"en","type":"article","venue":"Software Process Improvement and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Spice; Usability; Computer science; Empirical research; Software engineering; Dimension (graph theory); Process (computing); Reliability engineering; Systems engineering; Process management; Engineering; Mathematics","score_opus":0.05926722609592025,"score_gpt":0.3780309595307635,"score_spread":0.3187637334348432,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2040225071","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95384675,0.0010440266,0.00693696,0.0028918916,0.00020944487,0.017499367,0.0029567874,0.00017894177,0.014435884],"genre_scores_gemma":[0.94350463,0.0010131728,0.009500733,0.001768257,0.00018104809,0.036166698,0.0025612316,0.000115330135,0.005188916],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8177048,0.114537634,0.01793484,0.0045324783,0.040534355,0.0047558188],"domain_scores_gemma":[0.4229057,0.38966566,0.04489008,0.025399974,0.11309653,0.004042034],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.16026352,0.0011816431,0.002547371,0.003695173,0.0023885723,0.0039957436,0.002688153,0.0029687916,0.0037309474],"category_scores_gemma":[0.4746471,0.001045359,0.0015611582,0.004268822,0.0025522825,0.0045110174,0.0049256017,0.0044404175,0.001296239],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.04313991,0.04270242,0.13614601,0.009897168,0.0020506145,0.0017454916,0.16958867,0.009044793,0.007038371,0.008222815,0.035190735,0.53523314],"study_design_scores_gemma":[0.0142711345,0.13542722,0.5040972,0.0076606865,0.0023821492,0.0014176631,0.1368563,0.0080606295,0.046020254,0.009969946,0.13292995,0.00090682163],"about_ca_topic_score_codex":0.003355955,"about_ca_topic_score_gemma":0.0037202095,"teacher_disagreement_score":0.16026352,"about_ca_system_score_codex":0.0039528776,"about_ca_system_score_gemma":0.0059532532,"threshold_uncertainty_score":0.84756464},"labels":[],"label_agreement":null},{"id":"W2040411904","doi":"10.1007/s10664-014-9312-1","title":"Do topics make sense to managers and developers?","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Traceability; Latent Dirichlet allocation; Computer science; Documentation; Relevance (law); Requirements traceability; Perception; Topic model; Control (management); Software engineering; Requirements engineering; Requirements elicitation; BitTorrent tracker; Software; Data science; Information retrieval; Requirement; Artificial intelligence","score_opus":0.019201760776928176,"score_gpt":0.26889327776064786,"score_spread":0.2496915169837197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2040411904","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79928917,0.013769214,0.012090504,0.07471633,0.0034269597,0.00013910857,0.0010378245,0.00027192663,0.09525901],"genre_scores_gemma":[0.98537934,0.0035438214,0.002184164,0.0034348485,0.0009290126,0.00007738402,0.00036634918,0.00022233705,0.003862812],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9921967,0.003692909,0.00042102148,0.001020517,0.0015989419,0.0010699312],"domain_scores_gemma":[0.93719476,0.033595998,0.010698068,0.003962969,0.0064100586,0.008138141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008111592,0.0004849129,0.00061054726,0.0040911636,0.0026077083,0.012790399,0.0011270333,0.0025193563,0.0090684695],"category_scores_gemma":[0.07862103,0.0005182946,0.0005448722,0.005896791,0.004745812,0.022037033,0.004561534,0.002389319,0.0019104952],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007149133,0.00040128,0.409554,0.0013834756,0.000252793,0.0008177864,0.20411879,0.00028319898,0.0037858805,0.10394907,0.042358287,0.23238054],"study_design_scores_gemma":[0.00019297842,0.00038679512,0.3318449,0.0017625387,0.00041899094,0.00084679225,0.3458059,0.0010553626,0.002018507,0.16930121,0.14624043,0.00012564335],"about_ca_topic_score_codex":0.0024894467,"about_ca_topic_score_gemma":0.0029314095,"teacher_disagreement_score":0.012790399,"about_ca_system_score_codex":0.0022367279,"about_ca_system_score_gemma":0.003408029,"threshold_uncertainty_score":0.042898715},"labels":[],"label_agreement":null},{"id":"W2040952376","doi":"10.1145/1107656.1107667","title":"Traceability in viewpoint merging","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Mitacs","keywords":"Viewpoints; Merge (version control); Traceability; Computer science; Tracing; Set (abstract data type); Data mining; Information retrieval; Data science; Software engineering; Programming language","score_opus":0.016863276078136213,"score_gpt":0.27841917945601413,"score_spread":0.26155590337787793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2040952376","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032001638,0.00020833267,0.9918692,0.00023952272,0.00004056423,0.00011819853,0.000058299072,0.00044577869,0.0038199883],"genre_scores_gemma":[0.14013983,0.00070926646,0.85404253,0.00020368806,0.00009543718,0.0003631174,0.0006108708,0.00039512556,0.0034401454],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9770355,0.0110963285,0.0016419443,0.0035693897,0.0059242984,0.00073253375],"domain_scores_gemma":[0.95369923,0.021997452,0.0034123536,0.014898693,0.0053242925,0.00066788564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015511277,0.0016099939,0.0013081066,0.005225432,0.0029352477,0.0065885736,0.003720724,0.0028696477,0.004983702],"category_scores_gemma":[0.067082696,0.0018101525,0.0031798133,0.004940022,0.0066606193,0.017005129,0.009162416,0.005058354,0.0010808284],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014026403,0.000082907536,0.0023193823,0.00042907125,0.00010144141,0.0006575928,0.005297078,0.035534292,0.0036940454,0.7797258,0.0020327347,0.16998541],"study_design_scores_gemma":[0.00006474872,0.00013654359,0.0007775213,0.0003048631,0.00013583055,0.0005978694,0.0007826496,0.0852468,0.008696564,0.8577388,0.04541639,0.00010146845],"about_ca_topic_score_codex":0.006914488,"about_ca_topic_score_gemma":0.003974778,"teacher_disagreement_score":0.015511277,"about_ca_system_score_codex":0.0028026577,"about_ca_system_score_gemma":0.0035227253,"threshold_uncertainty_score":0.08203244},"labels":[],"label_agreement":null},{"id":"W2041075520","doi":"10.1109/icsm.2012.6405267","title":"An empirical study of build system migrations in practice: Case studies on KDE and the Linux kernel","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Computer science; Executable; Source code; Linux kernel; Process (computing); Deliverable; Software engineering; Codebase; Software; Software development; Code (set theory); World Wide Web; Operating system; Systems engineering; Engineering; Programming language; Set (abstract data type)","score_opus":0.057270343458378996,"score_gpt":0.408795564308688,"score_spread":0.351525220850309,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2041075520","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9975048,0.00010488449,0.0007811469,0.0004309633,0.0000037278696,0.000117877134,0.00003749829,0.000011963544,0.0010073042],"genre_scores_gemma":[0.9964664,0.00017346174,0.0024598988,0.00012778952,0.000004682574,0.00018553803,0.000074479925,0.000020827154,0.00048705484],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9759009,0.015825568,0.0014682811,0.0015870288,0.0034783192,0.0017399278],"domain_scores_gemma":[0.74756575,0.18894435,0.027654372,0.010729604,0.019122163,0.0059837485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028219549,0.0005306496,0.0006495806,0.0044368603,0.0040760217,0.00359468,0.002965058,0.0026539215,0.0017525394],"category_scores_gemma":[0.12746401,0.00097649265,0.00045578665,0.0044989754,0.0056949016,0.0077241603,0.0040578498,0.003374227,0.00043883425],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041964126,0.0050822604,0.34330517,0.0011226124,0.000119391036,0.0055683535,0.5673068,0.0023470242,0.0025523424,0.005579434,0.0033776849,0.063219264],"study_design_scores_gemma":[0.00014186405,0.0018900447,0.27570727,0.0007853638,0.000067442794,0.0017610022,0.69203544,0.009058693,0.002003488,0.0021033578,0.014302093,0.000143922],"about_ca_topic_score_codex":0.009987115,"about_ca_topic_score_gemma":0.019864183,"teacher_disagreement_score":0.028219549,"about_ca_system_score_codex":0.004675099,"about_ca_system_score_gemma":0.003658777,"threshold_uncertainty_score":0.14924103},"labels":[],"label_agreement":null},{"id":"W2041335584","doi":"10.1145/2499393.2499398","title":"An algorithmic approach to missing data problem in modeling human aspects in software development","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Code (set theory); Software; Software bug; Data modeling; Data mining; Machine learning; Artificial intelligence; Software engineering; Programming language","score_opus":0.07106634230404742,"score_gpt":0.3055231188397505,"score_spread":0.23445677653570307,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2041335584","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0074370885,0.00014412856,0.9910386,0.00058305444,0.000033864562,0.000091280104,0.00015898095,0.00016329942,0.0003496352],"genre_scores_gemma":[0.28322402,0.0003577463,0.71113497,0.00058368244,0.00030585632,0.0011842756,0.0010296134,0.0000978679,0.0020820536],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98650753,0.008569306,0.0006551035,0.0024728274,0.0013248828,0.00047033257],"domain_scores_gemma":[0.875997,0.1101783,0.0045794174,0.0048815752,0.0035179954,0.0008455807],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027293498,0.001747285,0.0026979488,0.0034707913,0.0017130442,0.0030712746,0.0057130777,0.0037128278,0.0039023405],"category_scores_gemma":[0.07856141,0.0018653088,0.0033430215,0.003690154,0.0033210055,0.0053386586,0.0046865977,0.0057882783,0.0005898451],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021536613,0.00040439403,0.01697318,0.0003270011,0.0005383571,0.00035514476,0.00062635983,0.82260394,0.00042031144,0.07269989,0.0023347563,0.08250128],"study_design_scores_gemma":[0.00003239841,0.000052869913,0.0006757936,0.00003522144,0.000041058553,0.00008333596,0.000042530206,0.93668926,0.00013877271,0.06172306,0.00046648041,0.00001926287],"about_ca_topic_score_codex":0.009828379,"about_ca_topic_score_gemma":0.010520297,"teacher_disagreement_score":0.027293498,"about_ca_system_score_codex":0.0022987362,"about_ca_system_score_gemma":0.0046902495,"threshold_uncertainty_score":0.14434355},"labels":[],"label_agreement":null},{"id":"W2041589121","doi":"10.5555/2820518.2820597","title":"The Firefox temporal defect dataset","year":2015,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software bug; Process (computing); Software; Plan (archaeology); Data mining; Data science; Geography","score_opus":0.030207281263329028,"score_gpt":0.28240052547346595,"score_spread":0.2521932442101369,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2041589121","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07422754,0.0013300122,0.0021242949,0.00051778374,0.00012702253,0.0002400476,0.9140127,0.0036336228,0.0037870803],"genre_scores_gemma":[0.020450784,0.0002843566,0.005572837,0.00011885214,0.000037306014,0.00020261767,0.9718784,0.000115445866,0.0013393163],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986951,0.000120478006,0.00016758201,0.00036249877,0.0004977438,0.00015652094],"domain_scores_gemma":[0.99709225,0.0007083876,0.0005041599,0.0005778714,0.0008228387,0.00029435626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012245564,0.0015040167,0.0006770855,0.0070736483,0.0007636699,0.00096230925,0.0020002269,0.0015365207,0.002736955],"category_scores_gemma":[0.0048390585,0.00041899376,0.0010749103,0.0048286514,0.00038262352,0.0012046547,0.00099546,0.0010335515,0.0039864425],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079582236,0.0009345521,0.08896516,0.0019552219,0.0003254124,0.0012932018,0.0004995013,0.005266598,0.005905402,0.0021573105,0.7999632,0.091938674],"study_design_scores_gemma":[0.00074469694,0.00069510005,0.2622463,0.0004957426,0.0002899564,0.0032359173,0.00076869014,0.025987111,0.0081309965,0.0033234123,0.6938828,0.00019928733],"about_ca_topic_score_codex":0.027174866,"about_ca_topic_score_gemma":0.051422946,"teacher_disagreement_score":0.027174866,"about_ca_system_score_codex":0.0014230446,"about_ca_system_score_gemma":0.0014793462,"threshold_uncertainty_score":0.0540334},"labels":[],"label_agreement":null},{"id":"W2041974913","doi":"10.1109/scam.2010.25","title":"Deriving Coupling Metrics from Call Graphs","year":2010,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Coupling (piping); Call graph; Theoretical computer science; Engineering","score_opus":0.028888926907979572,"score_gpt":0.2791611540567684,"score_spread":0.25027222714878883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2041974913","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23666508,0.00041936678,0.75213265,0.0001543523,0.000055791068,0.00019883386,0.00088048534,0.0052691684,0.0042242208],"genre_scores_gemma":[0.6537903,0.00029725643,0.34051907,0.000049524453,0.0000310334,0.0003035448,0.0025485228,0.0015511368,0.0009096731],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9938193,0.0017752207,0.0005473966,0.00066565984,0.0028938905,0.0002986154],"domain_scores_gemma":[0.9605137,0.022605127,0.0033418392,0.0060452325,0.007082403,0.00041163753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031552478,0.0015381991,0.00078266795,0.008450979,0.00062848383,0.0019786016,0.0012857413,0.00094359723,0.0016965299],"category_scores_gemma":[0.05423424,0.0005250866,0.00090034504,0.006236,0.0008310035,0.0033532428,0.0015945035,0.0010664032,0.00050572545],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031985185,0.00033685417,0.05731132,0.00084901485,0.0004033839,0.0002706188,0.00095039135,0.28410468,0.027967375,0.06524193,0.005694849,0.55654967],"study_design_scores_gemma":[0.000040900744,0.00021348518,0.019405378,0.000073688076,0.00013516743,0.00027914759,0.00023693725,0.8851033,0.026602414,0.06133635,0.0064769736,0.000096281205],"about_ca_topic_score_codex":0.0038021384,"about_ca_topic_score_gemma":0.003978667,"teacher_disagreement_score":0.008450979,"about_ca_system_score_codex":0.0013439461,"about_ca_system_score_gemma":0.0014816865,"threshold_uncertainty_score":0.016686738},"labels":[],"label_agreement":null},{"id":"W2042154985","doi":"10.1109/ase.2013.6693089","title":"Variability-aware performance prediction: A statistical learning approach","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":216,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Feature (linguistics); Machine learning; Artificial intelligence; Statistical learning; Data mining; Software; Correlation; Sample (material); Random forest; Mathematics","score_opus":0.012379645259124342,"score_gpt":0.22711479015022293,"score_spread":0.21473514489109857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2042154985","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027447766,0.00014287572,0.9700095,0.00020981215,0.000018420844,0.00006589131,0.00017645776,0.0013637336,0.0005656361],"genre_scores_gemma":[0.7720741,0.0002121494,0.22511756,0.00017427636,0.00011161144,0.00029212458,0.00086426095,0.00024888696,0.00090501155],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99601364,0.0013424694,0.0002926676,0.0010313774,0.001091204,0.00022873325],"domain_scores_gemma":[0.96819353,0.021979922,0.0030831776,0.003429986,0.0028869377,0.00042647985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064387918,0.0015746534,0.0016290925,0.0034152125,0.00054938963,0.0017296695,0.0025499675,0.0014332803,0.0008807048],"category_scores_gemma":[0.032950316,0.000687591,0.0011673094,0.002649077,0.00089740864,0.0029056934,0.0012964371,0.0029471025,0.00052086444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015127612,0.00035038454,0.022596361,0.00009622503,0.00027365494,0.00012644778,0.00016581855,0.775284,0.0026988275,0.005222602,0.0017439786,0.19129042],"study_design_scores_gemma":[0.0000043966247,0.000031061634,0.0010019325,0.000006769403,0.000013636816,0.000019438097,0.000010262361,0.99364364,0.0007579833,0.004342307,0.00015521431,0.000013314938],"about_ca_topic_score_codex":0.00466171,"about_ca_topic_score_gemma":0.004662867,"teacher_disagreement_score":0.0064387918,"about_ca_system_score_codex":0.0012864169,"about_ca_system_score_gemma":0.0012202087,"threshold_uncertainty_score":0.034052014},"labels":[],"label_agreement":null},{"id":"W2042459021","doi":"10.1145/2652524.2652586","title":"Effect of temporal collaboration network, maintenance activity, and experience on defect exposure","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); Toronto Metropolitan University","funders":"","keywords":"Temporality; Computer science; Context (archaeology); Plan (archaeology); Quality (philosophy); Software bug; Software","score_opus":0.00699742905476752,"score_gpt":0.2635620484217534,"score_spread":0.2565646193669859,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2042459021","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9958482,0.00019421444,0.002734983,0.000153585,0.00001466969,0.000021459407,0.00023977438,0.000042805448,0.00075025205],"genre_scores_gemma":[0.9989042,0.00005138618,0.0004988798,0.000009729354,0.000008316198,0.000017463022,0.00019759692,0.0000067824226,0.00030557977],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9976096,0.0011232374,0.00016740976,0.00052257156,0.0003466033,0.00023056589],"domain_scores_gemma":[0.88595617,0.09286801,0.0102619175,0.0031171988,0.0023163408,0.0054803635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056059975,0.0004833194,0.00033915514,0.0010701922,0.0003893384,0.0012753231,0.0008166354,0.0009909015,0.003989013],"category_scores_gemma":[0.051715482,0.00023184177,0.0008412064,0.00085213705,0.00044657668,0.0016917677,0.0010782661,0.0008435268,0.00036460866],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093573204,0.00038114982,0.97305715,0.000049497765,0.00034064133,0.00017632391,0.00030643705,0.010514218,0.0004822311,0.0002242029,0.0002952504,0.013237146],"study_design_scores_gemma":[0.000046627927,0.0013378001,0.90801436,0.00005039229,0.0004118834,0.00036543276,0.00071910914,0.08694758,0.000560633,0.0009719814,0.0005462053,0.000028087687],"about_ca_topic_score_codex":0.0069883093,"about_ca_topic_score_gemma":0.004899058,"teacher_disagreement_score":0.0069883093,"about_ca_system_score_codex":0.0006358645,"about_ca_system_score_gemma":0.00070574816,"threshold_uncertainty_score":0.029647708},"labels":[],"label_agreement":null},{"id":"W2042678439","doi":"10.1002/j.2334-5837.2013.tb03018.x","title":"9.5.2 New Opportunities for Architecture Measurement","year":2013,"lang":"en","type":"article","venue":"INCOSE International Symposium","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"Architecture; Nexus (standard); Suite; Computer science; Systems engineering; Plan (archaeology); Reference architecture; Systems architecture; Software engineering; Architecture framework; Process management; Engineering management; Engineering; Software architecture; Software; Embedded system","score_opus":0.05144478018099051,"score_gpt":0.26705292418420035,"score_spread":0.21560814400320982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2042678439","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046778478,0.0031006464,0.8696817,0.011710056,0.00058350613,0.00035526094,0.00090100674,0.003486952,0.063402295],"genre_scores_gemma":[0.35958883,0.0013381832,0.63194156,0.0005016313,0.00023550539,0.00047144465,0.0007657361,0.00035937264,0.0047976365],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9723823,0.017529607,0.0016657333,0.0020099161,0.005808344,0.00060412724],"domain_scores_gemma":[0.9544013,0.01695356,0.0040069222,0.013342558,0.0098294495,0.001466239],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024117945,0.001155953,0.0009789799,0.0052394355,0.0012192032,0.008901703,0.0021109832,0.0015507732,0.005847463],"category_scores_gemma":[0.03883819,0.0005834895,0.0010482036,0.0045126043,0.0028448638,0.011785737,0.0036777451,0.0034483578,0.0013894255],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014656647,0.0002886903,0.024234528,0.0005842067,0.000107474385,0.000078842444,0.0015770039,0.00832871,0.006757977,0.52700675,0.0099496655,0.42093965],"study_design_scores_gemma":[0.00007104011,0.00079964305,0.025160765,0.0012134142,0.000111032256,0.00041675888,0.002824918,0.112615645,0.017015595,0.6345256,0.20497483,0.00027065765],"about_ca_topic_score_codex":0.0019968438,"about_ca_topic_score_gemma":0.0018508654,"teacher_disagreement_score":0.024117945,"about_ca_system_score_codex":0.0021560837,"about_ca_system_score_gemma":0.0029691644,"threshold_uncertainty_score":0.12754941},"labels":[],"label_agreement":null},{"id":"W2043030767","doi":"10.1016/j.jss.2012.02.034","title":"Coding-error based defects in enterprise resource planning software: Prevention, discovery, elimination and mitigation","year":2012,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Coding (social sciences); Enterprise resource planning; Computer science; Software; Software engineering; Risk analysis (engineering); Knowledge management; Business; Operating system; Mathematics","score_opus":0.019357104095974048,"score_gpt":0.27485237154378017,"score_spread":0.25549526744780615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2043030767","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8456617,0.00032009685,0.14720064,0.00075408653,0.000060304457,0.00015942872,0.00025757923,0.002162831,0.0034233343],"genre_scores_gemma":[0.95613384,0.00011952253,0.04229106,0.00005317434,0.000007756542,0.000031786145,0.00021499861,0.00013625411,0.0010116515],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9946607,0.001517906,0.0005070773,0.00041146833,0.002613896,0.00028890077],"domain_scores_gemma":[0.8694269,0.08039465,0.016371043,0.015054572,0.01798026,0.0007726683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057015372,0.0004106594,0.0003238407,0.0027871113,0.0006577361,0.0016632306,0.0014690624,0.0010878759,0.0009470364],"category_scores_gemma":[0.07406184,0.00049209764,0.0004945832,0.002110741,0.0012598329,0.001962768,0.0011109636,0.0011382755,0.00023051421],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074081717,0.0007785658,0.30749413,0.000466757,0.00015149241,0.0019708984,0.0048894584,0.09867692,0.023290709,0.018343804,0.0027160812,0.5404804],"study_design_scores_gemma":[0.000094098075,0.0010682085,0.102501094,0.00042633392,0.00054988964,0.003959664,0.0031169415,0.744356,0.110789105,0.02791692,0.005029201,0.00019257507],"about_ca_topic_score_codex":0.009035799,"about_ca_topic_score_gemma":0.008389668,"teacher_disagreement_score":0.009035799,"about_ca_system_score_codex":0.0011442439,"about_ca_system_score_gemma":0.003433738,"threshold_uncertainty_score":0.030152977},"labels":[],"label_agreement":null},{"id":"W2043140003","doi":"10.1016/j.scico.2013.10.006","title":"SeByte: Scalable clone and similarity search for bytecode","year":2013,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of Saskatchewan; Queen's University","funders":"","keywords":"Computer science; Bytecode; Scalability; Java bytecode; Ranking (information retrieval); Data mining; Similarity (geometry); Nearest neighbor search; Jaccard index; Source code; Matching (statistics); Heuristic; Binary search algorithm; Information retrieval; Theoretical computer science; Java; Search algorithm; Artificial intelligence; Programming language; Database; Pattern recognition (psychology); Java applet; Java annotation","score_opus":0.0207267550115444,"score_gpt":0.28171656020423,"score_spread":0.2609898051926856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2043140003","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.061837558,0.001607503,0.7144258,0.00027952803,0.00030685906,0.000410419,0.0025323252,0.21405262,0.004547363],"genre_scores_gemma":[0.24321029,0.00052348396,0.7217562,0.00026883167,0.0001298876,0.00045677595,0.011024448,0.012789274,0.00984081],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976127,0.00027562177,0.00016191705,0.0005900507,0.0011318482,0.00022789282],"domain_scores_gemma":[0.9965239,0.0010920698,0.00026576882,0.0013932064,0.00053250085,0.00019250785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001113742,0.0021190867,0.0024420633,0.004351909,0.0013745048,0.0017411985,0.0049471622,0.0022511338,0.0126803],"category_scores_gemma":[0.007852188,0.0011865004,0.0017236368,0.0052514756,0.0010683734,0.0060237814,0.0049877614,0.0014737913,0.005087618],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012879524,0.00037320616,0.0049207318,0.00073907716,0.000362257,0.00046421544,0.0004969587,0.010090813,0.03836161,0.012142291,0.06712729,0.8636335],"study_design_scores_gemma":[0.00084913167,0.00079699943,0.003946506,0.0001306561,0.0002689686,0.0011571279,0.0005303647,0.80401385,0.08429966,0.059824206,0.043992974,0.00018945504],"about_ca_topic_score_codex":0.006743291,"about_ca_topic_score_gemma":0.011106877,"teacher_disagreement_score":0.0126803,"about_ca_system_score_codex":0.0010443332,"about_ca_system_score_gemma":0.002056101,"threshold_uncertainty_score":0.04241979},"labels":[],"label_agreement":null},{"id":"W2043445799","doi":"10.1049/iet-sen.2013.0165","title":"Analogy‐based effort estimation: a new method to discover set of analogies from dataset characteristics","year":2015,"lang":"en","type":"article","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"University of Jordan; National Aeronautics and Space Administration","keywords":"Computer science; Analogy; Medoid; Set (abstract data type); Cluster analysis; Data mining; Estimation; Machine learning; Artificial intelligence; Software; Engineering; Programming language","score_opus":0.061420497988652414,"score_gpt":0.3609502763795231,"score_spread":0.2995297783908707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2043445799","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019323189,0.00026871398,0.9771871,0.00016359496,0.000033884422,0.00015877654,0.0003474944,0.0010518361,0.0014653611],"genre_scores_gemma":[0.41174173,0.00039765143,0.5833023,0.00019669833,0.00009219359,0.0006526149,0.0014995735,0.0001674412,0.001949745],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.996581,0.0009318457,0.00033034358,0.0010536432,0.0009675036,0.00013558294],"domain_scores_gemma":[0.99061525,0.0052536116,0.001365396,0.0012567285,0.0013037001,0.00020540158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025591506,0.0013526672,0.0013151561,0.005976337,0.0006143794,0.0016481103,0.00206879,0.0012829748,0.0023140404],"category_scores_gemma":[0.022255981,0.0006059186,0.0015223338,0.004705694,0.00076745875,0.0044851387,0.0020365326,0.0018546565,0.00069626316],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045391815,0.00045090902,0.019573664,0.00051257765,0.00046259907,0.00023411207,0.0007126793,0.18860579,0.0067074643,0.020335536,0.0040467856,0.75790393],"study_design_scores_gemma":[0.000033852186,0.00020775042,0.0056816824,0.000052955816,0.000057000532,0.00018660155,0.000123223,0.9636274,0.0023858272,0.02377199,0.0038030357,0.00006863221],"about_ca_topic_score_codex":0.0027455871,"about_ca_topic_score_gemma":0.002672377,"teacher_disagreement_score":0.005976337,"about_ca_system_score_codex":0.0010366164,"about_ca_system_score_gemma":0.0011679772,"threshold_uncertainty_score":0.013534248},"labels":[],"label_agreement":null},{"id":"W2043547578","doi":"10.1109/iceccs.2010.5","title":"A Network Analysis of Stakeholders in Tool Visioning Process for Story Test Driven Development","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Agile software development; Centrality; Social network analysis; Computer science; Process (computing); User story; Knowledge management; Key (lock); Categorization; Process management; Social network (sociolinguistics); Software development; Software; Business; Software engineering; World Wide Web; Social media","score_opus":0.043896055926289264,"score_gpt":0.29485109012074284,"score_spread":0.2509550341944536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2043547578","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7680423,0.00010390579,0.21787268,0.00046697026,0.000014844231,0.00033776427,0.00042017127,0.00017883659,0.012562559],"genre_scores_gemma":[0.9533888,0.00005254568,0.044233534,0.000016886474,0.0000045768675,0.00015687272,0.00042651332,0.000027610995,0.00169268],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9976278,0.0014075856,0.000091506074,0.00030910384,0.0003945833,0.00016951712],"domain_scores_gemma":[0.9870156,0.009543373,0.0009862486,0.0004737367,0.0015719498,0.00040906164],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027841786,0.00029651463,0.00019683038,0.003415667,0.0014692156,0.0012290844,0.0005233862,0.0007389037,0.0023956024],"category_scores_gemma":[0.013358429,0.00022289912,0.0004641414,0.0019946736,0.00059008977,0.0029425777,0.0009831369,0.0004471027,0.00025656648],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018271665,0.00065985776,0.24736199,0.0008033505,0.0002659261,0.0037994399,0.0896858,0.07783707,0.043236103,0.12818734,0.0053158854,0.40102005],"study_design_scores_gemma":[0.000052966403,0.00044744162,0.13053226,0.00015274012,0.00019212006,0.0010484763,0.04059307,0.7506703,0.012533083,0.0432382,0.020401385,0.00013793119],"about_ca_topic_score_codex":0.0070330296,"about_ca_topic_score_gemma":0.007136992,"teacher_disagreement_score":0.0070330296,"about_ca_system_score_codex":0.001396651,"about_ca_system_score_gemma":0.00080768485,"threshold_uncertainty_score":0.014724314},"labels":[],"label_agreement":null},{"id":"W2044396529","doi":"10.1016/j.jss.2012.03.028","title":"Preserving knowledge in software projects","year":2012,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University; McGill University","funders":"Universiteit Gent; King Abdulaziz University; Polytechnique Montréal","keywords":"Scope (computer science); Computer science; Software engineering; Software development; Software; Exploit; Software mining; Knowledge management; Software construction; Computer security; Operating system; Programming language","score_opus":0.033872038305964074,"score_gpt":0.28261224862066975,"score_spread":0.24874021031470567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2044396529","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35049587,0.0011719979,0.5876931,0.006040743,0.00014258805,0.00017839068,0.00031454943,0.0007985812,0.05316413],"genre_scores_gemma":[0.96071875,0.00039827914,0.030084431,0.00016080204,0.00007623058,0.00006703008,0.000190313,0.00011585605,0.0081882],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9877422,0.005291416,0.0011790418,0.0014212237,0.003177565,0.0011885589],"domain_scores_gemma":[0.8873764,0.05892892,0.0054641785,0.03661098,0.007877966,0.0037414897],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010234087,0.0004666009,0.0010506685,0.002915506,0.0037934124,0.007928758,0.0028035827,0.0030624429,0.003520632],"category_scores_gemma":[0.069784895,0.0011447639,0.00142362,0.004866004,0.006944933,0.024140198,0.010893241,0.004017371,0.00061296666],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016792715,0.00015085097,0.0041541723,0.00012493707,0.000068316636,0.00036545368,0.0025687045,0.0073838253,0.00094564783,0.91956234,0.0010535349,0.063454285],"study_design_scores_gemma":[0.00001558777,0.00002837995,0.00031343344,0.000020482272,0.000034817072,0.00012445734,0.00035385755,0.005622594,0.00090272876,0.9896674,0.0029070647,0.000009091016],"about_ca_topic_score_codex":0.0025183894,"about_ca_topic_score_gemma":0.0017170552,"teacher_disagreement_score":0.010234087,"about_ca_system_score_codex":0.001847972,"about_ca_system_score_gemma":0.0033292181,"threshold_uncertainty_score":0.0541237},"labels":[],"label_agreement":null},{"id":"W2044527216","doi":"10.1109/saner.2015.7081813","title":"An observational study on API usage constraints and their documentation","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Documentation; Reuse; Computer science; Software documentation; Observational study; Constraint (computer-aided design); Software engineering; Software; Software development; Programming language; Software development process; Engineering","score_opus":0.18027514210196283,"score_gpt":0.3754529440937603,"score_spread":0.19517780199179746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2044527216","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9977664,0.000097551776,0.0010501309,0.0000938616,0.000006856804,0.00011722568,0.00025115913,0.000015262445,0.0006016726],"genre_scores_gemma":[0.99652004,0.00017663577,0.0019590498,0.000111110465,0.000015634434,0.00023244588,0.00043281855,0.000016044487,0.00053611025],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99406034,0.0028804257,0.00074515486,0.0006856814,0.0012495358,0.0003789606],"domain_scores_gemma":[0.920305,0.040979885,0.019179123,0.0055363113,0.010365706,0.0036339902],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056865,0.00032360858,0.0004391569,0.0017225468,0.0014136496,0.0008720692,0.00055081287,0.00085081253,0.0010269773],"category_scores_gemma":[0.048383217,0.0004984668,0.0003076633,0.000996441,0.0011634756,0.001250525,0.001721749,0.0017101291,0.00030223196],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028857504,0.0020210377,0.90648884,0.0002753033,0.000055662804,0.0012118093,0.05959518,0.00016178226,0.003986225,0.00034236323,0.0015866968,0.023986494],"study_design_scores_gemma":[0.0000582827,0.0024269493,0.9268292,0.00021880552,0.0000708328,0.0016575637,0.0478842,0.0013800312,0.0040450543,0.00053800974,0.014774282,0.0001167606],"about_ca_topic_score_codex":0.0042735613,"about_ca_topic_score_gemma":0.006171492,"teacher_disagreement_score":0.0056865,"about_ca_system_score_codex":0.00074971287,"about_ca_system_score_gemma":0.0013990286,"threshold_uncertainty_score":0.030073464},"labels":[],"label_agreement":null},{"id":"W2044807274","doi":"10.1007/s10664-010-9151-7","title":"Design evolution metrics for defect prediction in object oriented systems","year":2010,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Software metric; Eclipse; Data mining; Software quality; Quality (philosophy); Identification (biology); Software; Software system; Reliability engineering; Machine learning; Software engineering; Software development; Programming language; Engineering","score_opus":0.026845936819820337,"score_gpt":0.27604758517063077,"score_spread":0.24920164835081043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2044807274","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7136414,0.0017537595,0.28008926,0.00034148642,0.00006241484,0.00012693563,0.0006275161,0.0010898727,0.0022672801],"genre_scores_gemma":[0.9576934,0.00013739309,0.04125258,0.000017616814,0.000013890363,0.000052375162,0.0005364063,0.00004906827,0.00024729627],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9949469,0.0020686311,0.00049699633,0.0003032809,0.002017994,0.00016630343],"domain_scores_gemma":[0.94753146,0.03561074,0.006036671,0.0036512038,0.0063079596,0.0008619447],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006024601,0.000837305,0.0007142458,0.0065585733,0.00035795808,0.0009885108,0.0008169054,0.000871418,0.0007256448],"category_scores_gemma":[0.051922303,0.00031655942,0.00046581565,0.0026513012,0.00049301086,0.0021722037,0.0008167688,0.0008065454,0.00013949182],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044989784,0.0005311737,0.26521513,0.00033824253,0.00033295923,0.00008933922,0.0003206797,0.23625277,0.006020706,0.0064942944,0.0023338534,0.48162097],"study_design_scores_gemma":[0.00003282504,0.0006095235,0.042155113,0.000052727904,0.00009597361,0.00011720542,0.000070511654,0.94567645,0.0036427355,0.0069818,0.00053841044,0.00002670314],"about_ca_topic_score_codex":0.0026465973,"about_ca_topic_score_gemma":0.002880551,"teacher_disagreement_score":0.0065585733,"about_ca_system_score_codex":0.0010087567,"about_ca_system_score_gemma":0.0007902145,"threshold_uncertainty_score":0.031861484},"labels":[],"label_agreement":null},{"id":"W2045336717","doi":"10.1109/icsme.2014.31","title":"An Exploratory Study on Self-Admitted Technical Debt","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":282,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Technical debt; Debt; Workaround; Computer science; Software; Business; Software development; Finance; Operating system","score_opus":0.021141479268621637,"score_gpt":0.29018461736761114,"score_spread":0.2690431380989895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2045336717","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9990521,0.000026589285,0.00023188103,0.00006538272,0.0000018364881,0.000039129332,0.000037350786,0.0000052563028,0.0005404636],"genre_scores_gemma":[0.99802446,0.00009963372,0.00072533655,0.000109374276,0.0000061170276,0.000105928746,0.00011504686,0.0000099718645,0.0008041119],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9933102,0.0032764964,0.00058088114,0.000513077,0.0017217983,0.0005975112],"domain_scores_gemma":[0.88721746,0.069988124,0.025426483,0.0030503157,0.010449604,0.003868006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008240734,0.0002857366,0.0003433619,0.0024992514,0.0014605827,0.0020772486,0.0009678487,0.0008376114,0.0012456154],"category_scores_gemma":[0.057624552,0.00043938134,0.0002200931,0.0020495923,0.0013269249,0.0027191602,0.0019451698,0.0013351439,0.00028873142],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015927605,0.00068337517,0.58170414,0.00035883693,0.000023425519,0.0022422702,0.3842608,0.0001213097,0.0034577071,0.0009188823,0.0009606331,0.025109347],"study_design_scores_gemma":[0.000016102798,0.0006867776,0.5382374,0.0002776824,0.00001596626,0.0014950458,0.44684216,0.0012029482,0.0012256254,0.00048613656,0.009453303,0.000060869876],"about_ca_topic_score_codex":0.0028423949,"about_ca_topic_score_gemma":0.004778357,"teacher_disagreement_score":0.008240734,"about_ca_system_score_codex":0.001480499,"about_ca_system_score_gemma":0.001620934,"threshold_uncertainty_score":0.043581665},"labels":[],"label_agreement":null},{"id":"W2045352263","doi":"10.1145/965660.965676","title":"Refactoring to aspects","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Eclipse; Programming language; Software engineering; Software evolution; Software maintenance; Software; Restructuring; Extreme programming; Software system; Software development; Software construction; Software development process","score_opus":0.02508164165654799,"score_gpt":0.2770406081750777,"score_spread":0.2519589665185297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2045352263","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024116619,0.0020558322,0.93862164,0.001158746,0.000514327,0.00039326266,0.000260211,0.006495762,0.026383612],"genre_scores_gemma":[0.2177903,0.003271183,0.73399055,0.0017552314,0.00040587076,0.0003399598,0.0016087927,0.004188322,0.036649708],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9967805,0.00096953276,0.00027507238,0.0005495219,0.0011273994,0.00029784767],"domain_scores_gemma":[0.9878364,0.0030519436,0.00059646455,0.0060947957,0.002143787,0.00027655796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032036009,0.0009949537,0.0006316882,0.0014784519,0.00091772125,0.0018292468,0.001934529,0.0013565297,0.0047337757],"category_scores_gemma":[0.016617674,0.0006631008,0.0014863188,0.0010981143,0.0017752919,0.003159396,0.0036755162,0.0028246266,0.002404869],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023206635,0.00023851925,0.004832524,0.0014110183,0.00015648428,0.0018337969,0.0046200636,0.009829804,0.049913734,0.2939784,0.020979809,0.6119738],"study_design_scores_gemma":[0.00011504397,0.00027692658,0.0020074572,0.00063554617,0.00027087118,0.0020967415,0.00068364834,0.031447425,0.07557841,0.20682333,0.6799653,0.000099298275],"about_ca_topic_score_codex":0.0016475713,"about_ca_topic_score_gemma":0.0021071725,"teacher_disagreement_score":0.0047337757,"about_ca_system_score_codex":0.00078679953,"about_ca_system_score_gemma":0.0013109896,"threshold_uncertainty_score":0.016942441},"labels":[],"label_agreement":null},{"id":"W2045366549","doi":"10.1109/esem.2009.5316047","title":"A detailed examination of the correlation between imports and failure-proneness of software components","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Blame; Component (thermodynamics); Computer science; Correlation; Type (biology); Reliability engineering; Engineering; Mathematics; Psychology; Social psychology; Geology","score_opus":0.016344318269213867,"score_gpt":0.2378630335803394,"score_spread":0.22151871531112552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2045366549","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9916911,0.00033644142,0.005665021,0.00014286928,0.0000070939323,0.000015539352,0.000268011,0.00008941495,0.0017845551],"genre_scores_gemma":[0.99763525,0.000080855694,0.00166315,0.000016056665,0.000008553777,0.000006470155,0.00022816457,0.000026398116,0.00033508544],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99158996,0.0024161777,0.0009445538,0.0012371155,0.0032857182,0.0005265211],"domain_scores_gemma":[0.53927416,0.3483134,0.06695977,0.025886921,0.017527616,0.0020381007],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0098996665,0.00053082063,0.00060418714,0.003286961,0.0005571659,0.0015324367,0.0008204221,0.000684607,0.003945261],"category_scores_gemma":[0.0973679,0.0005729492,0.001067351,0.0037054548,0.0014715091,0.0030815715,0.0012541757,0.0020028262,0.00041450773],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013436342,0.00004275918,0.98475236,0.00007144758,0.00016274185,0.00015528817,0.000683754,0.001562473,0.0016434777,0.000639999,0.0001471966,0.01000411],"study_design_scores_gemma":[0.0000024769918,0.00010672039,0.9947389,0.000017236121,0.00004807806,0.00049776275,0.0003729267,0.0021758494,0.0010381471,0.0006143071,0.00036868113,0.000018981118],"about_ca_topic_score_codex":0.0022129705,"about_ca_topic_score_gemma":0.004008968,"teacher_disagreement_score":0.0098996665,"about_ca_system_score_codex":0.0005894492,"about_ca_system_score_gemma":0.0006516295,"threshold_uncertainty_score":0.05235505},"labels":[],"label_agreement":null},{"id":"W2046040805","doi":"10.1145/2735399.2735419","title":"Technical Debt","year":2015,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Maturity (psychological); Debt; Capability Maturity Model; Software; Software engineering; Computer science; Engineering management; Engineering; Business; Software development; Political science; Finance","score_opus":0.031984929105533105,"score_gpt":0.26965765861627405,"score_spread":0.23767272951074095,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046040805","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06071372,0.0076587936,0.017866485,0.048483297,0.004723429,0.0003657834,0.0043148436,0.0007252984,0.8551482],"genre_scores_gemma":[0.5185717,0.008614138,0.006249187,0.017901601,0.0031481227,0.00042045763,0.009281796,0.00085194124,0.43496108],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9925061,0.0013242477,0.00060590496,0.0007726272,0.0035109152,0.0012802639],"domain_scores_gemma":[0.9689132,0.0053121457,0.004725646,0.0041148025,0.011400929,0.0055333613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006916165,0.0004181066,0.00047435187,0.0026255047,0.003457982,0.008069057,0.0014792645,0.0016875763,0.089225695],"category_scores_gemma":[0.040951524,0.0002652796,0.00044731863,0.0043109795,0.001077188,0.007958448,0.0057450742,0.0036290404,0.017583273],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013825165,0.00020069117,0.02471697,0.00044293216,0.000040142888,0.0013955577,0.0055199685,0.00060407066,0.0008586424,0.27062577,0.41083297,0.28462404],"study_design_scores_gemma":[0.000012964037,0.000053095195,0.008036932,0.00031061596,0.000013641606,0.0009215143,0.0020837698,0.00025736043,0.00030436157,0.030709805,0.95727175,0.000024207033],"about_ca_topic_score_codex":0.002575396,"about_ca_topic_score_gemma":0.0025333432,"teacher_disagreement_score":0.089225695,"about_ca_system_score_codex":0.003438447,"about_ca_system_score_gemma":0.0038955607,"threshold_uncertainty_score":0.2984897},"labels":[],"label_agreement":null},{"id":"W2046157646","doi":"10.1007/s10664-015-9377-5","title":"An empirical study of software release notes","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Empirical research; Computer science; Software; Software engineering; Programming language; Mathematics; Statistics","score_opus":0.05553808588465858,"score_gpt":0.33739025458356287,"score_spread":0.2818521686989043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046157646","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9962166,0.00014002995,0.00039026293,0.00014098483,0.0000052673813,0.000029145434,0.00015134005,0.000012166576,0.002914267],"genre_scores_gemma":[0.997127,0.00014870186,0.00041150747,0.00005258419,0.000010960087,0.000020457595,0.00036821322,0.000011430103,0.0018491136],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9963671,0.0014713402,0.00030052432,0.00032417045,0.0012669765,0.00026991492],"domain_scores_gemma":[0.8280749,0.12994438,0.02068837,0.005924765,0.011597888,0.003769697],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004191337,0.00023331062,0.00019826245,0.0023960355,0.0011948597,0.0017062909,0.0010298648,0.0006020519,0.005437764],"category_scores_gemma":[0.06375992,0.00027278898,0.00017046431,0.00268573,0.0011237711,0.0028568585,0.001039571,0.0014197918,0.0012073813],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005237376,0.003656942,0.9043496,0.00034440312,0.000050926206,0.00087235303,0.020885907,0.00037578196,0.0022400983,0.0043486585,0.0031625356,0.059189096],"study_design_scores_gemma":[0.000040663792,0.00090906327,0.95967484,0.00015410886,0.000043060336,0.00060586876,0.027436469,0.0016948346,0.0017003284,0.00069371844,0.007013414,0.000033562654],"about_ca_topic_score_codex":0.0047257757,"about_ca_topic_score_gemma":0.0074738995,"teacher_disagreement_score":0.99580866,"about_ca_system_score_codex":0.0009963281,"about_ca_system_score_gemma":0.0011620091,"threshold_uncertainty_score":0.022166193},"labels":[],"label_agreement":null},{"id":"W2046253855","doi":"10.1145/2491411.2494587","title":"Code fragment summarization","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Automatic summarization; Code (set theory); Fragment (logic); Oracle; Code review; Programming language; Source code; Information retrieval; Software; Presentation (obstetrics); Data mining; Static program analysis; Software development","score_opus":0.013442418359849984,"score_gpt":0.2420684419017109,"score_spread":0.2286260235418609,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046253855","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10742417,0.0045008576,0.8117605,0.0014559391,0.00054775405,0.0013798577,0.01642031,0.045103066,0.011407518],"genre_scores_gemma":[0.27132273,0.0016510524,0.6583597,0.00033598454,0.00038008296,0.00051574304,0.052695345,0.003110032,0.011629341],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998214,0.00033018886,0.00019885617,0.00049898436,0.00064953556,0.00010847357],"domain_scores_gemma":[0.990838,0.0024111574,0.00096133654,0.0015986478,0.0039884984,0.00020241832],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015278363,0.0015273868,0.0012312516,0.0073645776,0.00079756364,0.0018978205,0.0014881422,0.0009114662,0.0070425924],"category_scores_gemma":[0.015316649,0.00026558634,0.0007813791,0.0044577,0.00034781685,0.0019859748,0.0013725988,0.00089531014,0.0036563238],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052493205,0.00010259343,0.0037450162,0.0015028444,0.00011577827,0.00032179974,0.00071813434,0.008197704,0.04046784,0.0032828425,0.049331468,0.8916891],"study_design_scores_gemma":[0.00023355901,0.0012447889,0.021848027,0.00063867774,0.0006337697,0.0027021805,0.0019351307,0.4525038,0.2386996,0.030572314,0.2486928,0.0002952502],"about_ca_topic_score_codex":0.003451555,"about_ca_topic_score_gemma":0.0036816818,"teacher_disagreement_score":0.0073645776,"about_ca_system_score_codex":0.00072767655,"about_ca_system_score_gemma":0.0010660163,"threshold_uncertainty_score":0.023559809},"labels":[],"label_agreement":null},{"id":"W2046347832","doi":"10.1145/2593801.2593803","title":"A mapping study on bayesian networks for software quality prediction","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Tekes","keywords":"Bayesian network; Computer science; Categorical variable; Machine learning; Artificial intelligence; Dependency (UML); Inference; Data mining; Software; Quality (philosophy); Software quality; Software development","score_opus":0.03577368909895908,"score_gpt":0.30628662842953874,"score_spread":0.2705129393305797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046347832","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29874206,0.015635116,0.670211,0.0016838369,0.00012595777,0.0007936709,0.00090313895,0.00020789636,0.01169735],"genre_scores_gemma":[0.7769607,0.007904327,0.21190353,0.00032794438,0.0000920821,0.00078701857,0.0008166043,0.000055218523,0.001152602],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9803115,0.015626961,0.0006121736,0.001398029,0.001884556,0.00016685616],"domain_scores_gemma":[0.8382546,0.14919746,0.0029784795,0.0034885106,0.00579384,0.00028716843],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.019382695,0.0009308046,0.00081720634,0.0061697196,0.00065001193,0.0015980408,0.0008363047,0.0009195624,0.0035954777],"category_scores_gemma":[0.121199585,0.0005005823,0.0019106151,0.0075039687,0.0009977653,0.0056288503,0.0012325271,0.0011304079,0.00032295735],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005866202,0.0007207608,0.09565971,0.005358932,0.0018429457,0.0006894208,0.004179693,0.072629906,0.0032552413,0.081357636,0.001890966,0.7318282],"study_design_scores_gemma":[0.0003338474,0.0028208392,0.08574583,0.0058310516,0.0032025077,0.0018345873,0.005135423,0.51735693,0.008660243,0.34450862,0.024342177,0.00022793331],"about_ca_topic_score_codex":0.0027642432,"about_ca_topic_score_gemma":0.0019419611,"teacher_disagreement_score":0.99383026,"about_ca_system_score_codex":0.001195725,"about_ca_system_score_gemma":0.0017044205,"threshold_uncertainty_score":0.10250676},"labels":[],"label_agreement":null},{"id":"W2046633649","doi":"10.1111/j.1467-8640.2006.00283.x","title":"VISIO‐SPATIAL CASE‐BASED REASONING: A CASE STUDY IN PREDICTION OF PROTEIN STRUCTURE","year":2006,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Carleton University","funders":"","keywords":"Similarity (geometry); Artificial intelligence; Computer science; Metric (unit); Case-based reasoning; Pattern recognition (psychology); Data mining; Image (mathematics)","score_opus":0.025771312733067,"score_gpt":0.3037859787183627,"score_spread":0.27801466598529573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046633649","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4223762,0.0014278493,0.542561,0.0061844746,0.00009643944,0.0004929557,0.0007046176,0.0009082529,0.025248257],"genre_scores_gemma":[0.68452775,0.0006190786,0.31108707,0.00019136013,0.000024283185,0.00016747201,0.00043512988,0.000103600694,0.0028442056],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969658,0.0017761482,0.00012690121,0.00020460972,0.00078937545,0.00013712345],"domain_scores_gemma":[0.98821473,0.009540795,0.00046022536,0.00094865024,0.00055173057,0.00028386022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040192185,0.00066637737,0.00043436544,0.0013141212,0.0021771227,0.0026443515,0.0022408713,0.0034597407,0.0024857149],"category_scores_gemma":[0.015051165,0.0004120482,0.0009716196,0.0017552276,0.0031867449,0.0036815682,0.0017922027,0.0020995932,0.00043955605],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013423028,0.002432421,0.03961654,0.0010869215,0.0002476715,0.040146217,0.027812118,0.18655647,0.014984894,0.3363795,0.02042408,0.32897085],"study_design_scores_gemma":[0.00035893387,0.00052673137,0.006724486,0.00042857908,0.00013993004,0.015067932,0.008859299,0.7164521,0.025765603,0.16232108,0.0631941,0.00016135118],"about_ca_topic_score_codex":0.0090644,"about_ca_topic_score_gemma":0.010517964,"teacher_disagreement_score":0.0090644,"about_ca_system_score_codex":0.0015596027,"about_ca_system_score_gemma":0.0010410973,"threshold_uncertainty_score":0.02125591},"labels":[],"label_agreement":null},{"id":"W2046816618","doi":"10.1115/detc2010-28773","title":"Estimation of the Scope of Change Propagation in Object-Oriented Programming","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Scope (computer science); Computer science; Dependency (UML); Representation (politics); Object (grammar); Context (archaeology); Object-oriented programming; Core (optical fiber); Theoretical computer science; Software engineering; Programming language; Artificial intelligence","score_opus":0.026335403537414013,"score_gpt":0.2837457319705648,"score_spread":0.25741032843315076,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046816618","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3453002,0.0003853448,0.65268016,0.000054601576,0.000006985929,0.00010811432,0.000048482387,0.00025619936,0.0011599123],"genre_scores_gemma":[0.8486364,0.0001870292,0.15080129,0.0000079936035,0.000013295859,0.00006230785,0.00010277032,0.000022995126,0.00016593764],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99562585,0.0020540394,0.00032978595,0.00051276083,0.00132099,0.00015647743],"domain_scores_gemma":[0.91697544,0.06742007,0.0081746075,0.003111481,0.0035759655,0.00074249785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076035587,0.0007852341,0.0006139875,0.004665583,0.00051167177,0.0010713802,0.0009652693,0.0010456517,0.00034005713],"category_scores_gemma":[0.054970574,0.0006922022,0.00058639416,0.0016193581,0.00084572437,0.0029672,0.0012470328,0.0006590331,0.00007710271],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004170112,0.00026058516,0.14981626,0.00038670577,0.0001732111,0.00059213256,0.0019076577,0.568337,0.012669913,0.015972402,0.00024915236,0.24921803],"study_design_scores_gemma":[0.000012524345,0.00022342023,0.02432675,0.000042296688,0.000040645147,0.00041932784,0.00032212085,0.95667046,0.0063444306,0.011167172,0.0003843366,0.000046541576],"about_ca_topic_score_codex":0.0029100373,"about_ca_topic_score_gemma":0.0022902666,"teacher_disagreement_score":0.0076035587,"about_ca_system_score_codex":0.0006428469,"about_ca_system_score_gemma":0.0007206478,"threshold_uncertainty_score":0.040211916},"labels":[],"label_agreement":null},{"id":"W2047000830","doi":"10.1109/vissof.2007.4290720","title":"Visual Analysis of Azureus using VERSO","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Task (project management); Computer science; Artificial intelligence; Engineering","score_opus":0.026610504259394166,"score_gpt":0.3423521174109593,"score_spread":0.3157416131515651,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2047000830","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10743284,0.002318453,0.6296737,0.0022746937,0.0014945279,0.00057542697,0.014592414,0.041903928,0.19973397],"genre_scores_gemma":[0.44949508,0.0021320435,0.45444882,0.0005833265,0.00050776027,0.00038759474,0.016131086,0.012297438,0.06401678],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99980134,0.00004934566,0.00001465627,0.00005436804,0.00005395714,0.0000263581],"domain_scores_gemma":[0.99931884,0.00031980267,0.000038846785,0.00009777049,0.00016369618,0.00006113526],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004996034,0.00094376574,0.0004152271,0.0019581711,0.0007375618,0.0024584013,0.00046663205,0.00051466684,0.03910017],"category_scores_gemma":[0.002185388,0.00023027562,0.0006438316,0.00089587166,0.0005428386,0.002248292,0.0010454814,0.0008758061,0.007132871],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001299532,0.00011924306,0.0052573057,0.0025303422,0.00013823654,0.0028573882,0.010074326,0.011447134,0.15750493,0.052475993,0.13387385,0.6224218],"study_design_scores_gemma":[0.00010770968,0.00028960226,0.016805982,0.00095932867,0.00008845199,0.0018715601,0.005634197,0.097207084,0.054766018,0.0455286,0.7765041,0.00023740882],"about_ca_topic_score_codex":0.0028610933,"about_ca_topic_score_gemma":0.0035179816,"teacher_disagreement_score":0.03910017,"about_ca_system_score_codex":0.00046364742,"about_ca_system_score_gemma":0.000300437,"threshold_uncertainty_score":0.13080311},"labels":[],"label_agreement":null},{"id":"W2047340539","doi":"","title":"Near-miss Model Clone Detection for Simulink Models","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Leverage (statistics); clone (Java method); Granularity; Source code; Detector; Identification (biology); Block (permutation group theory); Data mining; Programming language; Artificial intelligence","score_opus":0.032615082029270305,"score_gpt":0.2645018033448053,"score_spread":0.23188672131553498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2047340539","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.080041945,0.00010961982,0.9070227,0.00010254473,0.000027514128,0.000082515115,0.00017198251,0.01153394,0.00090713723],"genre_scores_gemma":[0.5420365,0.000101397534,0.45387638,0.000083100706,0.00000912054,0.00008740497,0.00052358577,0.0013629438,0.001919636],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997382,0.00051504263,0.00016335477,0.00038584444,0.0014480224,0.00010566678],"domain_scores_gemma":[0.9875557,0.0056535313,0.0019090407,0.0024179001,0.0023161655,0.00014774801],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012684042,0.00083701056,0.00057725376,0.0016055858,0.00043126053,0.0012908636,0.0010809225,0.0009419067,0.00162328],"category_scores_gemma":[0.01398642,0.00040166956,0.0008425504,0.00083687145,0.00070587324,0.0018264488,0.0013053517,0.0012876838,0.0005789257],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011992741,0.0003313349,0.049831916,0.0008628233,0.00025488643,0.0023481047,0.0023376883,0.21591589,0.23910469,0.038061846,0.0043140203,0.44543743],"study_design_scores_gemma":[0.000018692012,0.00018464179,0.001889462,0.00005698505,0.000046140332,0.00049599103,0.0001689389,0.81017447,0.17209971,0.0082542235,0.0065742037,0.000036607584],"about_ca_topic_score_codex":0.0024021307,"about_ca_topic_score_gemma":0.0034910054,"teacher_disagreement_score":0.0024021307,"about_ca_system_score_codex":0.001060596,"about_ca_system_score_gemma":0.00081612467,"threshold_uncertainty_score":0.0076951385},"labels":[],"label_agreement":null},{"id":"W2047550520","doi":"10.1109/wcre.2013.6671307","title":"Improving SOA antipatterns detection in Service Based Systems by mining execution traces","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Maintainability; Quality of service; Service-oriented architecture; Service (business); OASIS SOA Reference Model; Reusability; Data mining; Software engineering; Distributed computing; Software; Web service; Operating system; Computer network; Programming language","score_opus":0.011726348508564795,"score_gpt":0.22501544505542134,"score_spread":0.21328909654685654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2047550520","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3129853,0.00050339056,0.67240095,0.00033198108,0.000058153324,0.00026278372,0.00081678666,0.011170253,0.0014704159],"genre_scores_gemma":[0.60859096,0.00031851197,0.38741812,0.000080092184,0.00002334956,0.0001326687,0.0019731445,0.00027993735,0.0011831718],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99665123,0.00073812983,0.00039415795,0.0006338934,0.0014182794,0.00016435869],"domain_scores_gemma":[0.98625666,0.0069351094,0.0019183598,0.0016768125,0.0028579035,0.0003550574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023908175,0.001110742,0.0008322353,0.005401733,0.00053478614,0.001399503,0.0010750971,0.0007841973,0.0006017029],"category_scores_gemma":[0.014090053,0.0003711009,0.0009944552,0.0024925524,0.00039712942,0.0013983626,0.0012765615,0.0008040239,0.00059021235],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040360063,0.0006283696,0.124255985,0.00069989247,0.00027731102,0.0008127707,0.0010860822,0.028221752,0.06207628,0.002135367,0.002307589,0.777095],"study_design_scores_gemma":[0.000045296252,0.000346273,0.03170267,0.000093117545,0.00012976189,0.0010365396,0.00050228165,0.8875719,0.0665862,0.006810277,0.005112286,0.00006352063],"about_ca_topic_score_codex":0.004844235,"about_ca_topic_score_gemma":0.00689265,"teacher_disagreement_score":0.005401733,"about_ca_system_score_codex":0.0005063721,"about_ca_system_score_gemma":0.0014206107,"threshold_uncertainty_score":0.012643993},"labels":[],"label_agreement":null},{"id":"W2048194532","doi":"10.1109/sera.2011.45","title":"Quantifying the Impact of Different Non-functional Requirements and Problem Domains on Software Effort Estimation","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Estimation; Software; Reuse; Domain (mathematical analysis); Data mining; Machine learning; Feature (linguistics); Software metric; Selection (genetic algorithm); Software development; Artificial intelligence; Software quality; Mathematics; Engineering; Systems engineering","score_opus":0.07475897787646796,"score_gpt":0.315573614894095,"score_spread":0.24081463701762706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2048194532","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8587798,0.00023704318,0.13972494,0.00010056576,0.0000116129895,0.00007531496,0.00014204,0.0002916045,0.00063716643],"genre_scores_gemma":[0.95431966,0.00007309879,0.044997234,0.00002260999,0.00000861409,0.000063681335,0.00033024666,0.000035491692,0.00014935395],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9933814,0.003382189,0.0005356527,0.0010371867,0.0014124901,0.00025112648],"domain_scores_gemma":[0.8534707,0.12709409,0.0073157107,0.007241321,0.004418353,0.0004597944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00836679,0.0008628117,0.0007124735,0.0018304678,0.0002758458,0.00092901126,0.00062922644,0.0008271899,0.00028671845],"category_scores_gemma":[0.0789228,0.00039158124,0.00073935516,0.0015776084,0.00048700543,0.0019434236,0.0009827614,0.0008730623,0.00009730228],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004845443,0.0007940397,0.20964858,0.00029109945,0.00044573165,0.00014835993,0.0006391911,0.501518,0.013813272,0.0012782951,0.00048646095,0.27045244],"study_design_scores_gemma":[0.000024777555,0.0005120745,0.08853192,0.00002580348,0.000086481086,0.0001512408,0.00012949569,0.89678705,0.011666178,0.0017040307,0.00034235264,0.000038659775],"about_ca_topic_score_codex":0.0018554899,"about_ca_topic_score_gemma":0.00346203,"teacher_disagreement_score":0.00836679,"about_ca_system_score_codex":0.0006023036,"about_ca_system_score_gemma":0.0005960998,"threshold_uncertainty_score":0.044248343},"labels":[],"label_agreement":null},{"id":"W2048909176","doi":"10.1109/scam.2014.33","title":"ACUA: API Change and Usage Auditor","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Audit; Computer science; Accounting; Business","score_opus":0.0334765347362818,"score_gpt":0.2586926426319207,"score_spread":0.2252161078956389,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2048909176","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07925234,0.0008511703,0.28028828,0.0006415498,0.0004218248,0.0014096601,0.008935521,0.61941063,0.008789052],"genre_scores_gemma":[0.6493878,0.000590181,0.30484846,0.000683248,0.00025598207,0.0019583656,0.012311971,0.019122425,0.01084162],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9921251,0.0019026283,0.0009107077,0.0017064326,0.0028604143,0.00049484504],"domain_scores_gemma":[0.94807035,0.020701902,0.007240401,0.014163611,0.008346304,0.0014773381],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005194331,0.0022432778,0.0010708873,0.008802317,0.00069853023,0.0023041351,0.0024538112,0.0010814056,0.0058106165],"category_scores_gemma":[0.03834077,0.0011007431,0.0008035256,0.0026534128,0.00076951436,0.0032085045,0.0021647916,0.0019420369,0.0028046465],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016748161,0.0012072632,0.09792263,0.0014989006,0.0003131267,0.0009925829,0.0022292666,0.013104982,0.022319123,0.0064773127,0.09576698,0.75649303],"study_design_scores_gemma":[0.00042522085,0.00097514637,0.09132982,0.00068820996,0.00038392906,0.0019657505,0.0009791077,0.5245865,0.17066635,0.012013281,0.19526239,0.0007243508],"about_ca_topic_score_codex":0.0053047324,"about_ca_topic_score_gemma":0.0031221418,"teacher_disagreement_score":0.008802317,"about_ca_system_score_codex":0.0009067687,"about_ca_system_score_gemma":0.0023432293,"threshold_uncertainty_score":0.027470589},"labels":[],"label_agreement":null},{"id":"W2049055220","doi":"10.1109/icsm.2010.5609739","title":"Playing with refactoring: Identifying extract class opportunities through game theory","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code refactoring; Computer science; Cohesion (chemistry); Game theory; Task (project management); Class (philosophy); Management science; Software engineering; Quality (philosophy); Algorithmic game theory; Software; Artificial intelligence; Sequential game; Systems engineering; Engineering","score_opus":0.09029389536732674,"score_gpt":0.31051315937731266,"score_spread":0.22021926400998593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049055220","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10965665,0.00019283,0.8842593,0.00074046897,0.000029573079,0.00044280407,0.00007979585,0.00054497604,0.0040536434],"genre_scores_gemma":[0.5079276,0.00015354445,0.4899995,0.00013534918,0.000014483996,0.00023777301,0.00009975181,0.000067827954,0.0013641927],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99710053,0.0015249365,0.0001378712,0.0003841833,0.0006476346,0.00020472598],"domain_scores_gemma":[0.98794544,0.009521525,0.0006856084,0.0005578711,0.000956971,0.0003326659],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041267443,0.0011653106,0.0013031388,0.0019161069,0.00091259286,0.0023403552,0.0020602127,0.0016085061,0.001931939],"category_scores_gemma":[0.01845154,0.0005764282,0.0012401822,0.00087062374,0.0018235883,0.0037755303,0.0015520008,0.0016277907,0.00028972354],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00097597315,0.0018083041,0.03359791,0.0007505688,0.0003494111,0.00062715303,0.0044685034,0.3302412,0.02221347,0.1517597,0.0040163524,0.4491914],"study_design_scores_gemma":[0.00005233961,0.000182051,0.0015427191,0.00003253506,0.000044586126,0.00009613452,0.00027969448,0.9604895,0.0027965647,0.032842103,0.0016027365,0.000038996335],"about_ca_topic_score_codex":0.0074651153,"about_ca_topic_score_gemma":0.009560687,"teacher_disagreement_score":0.0074651153,"about_ca_system_score_codex":0.0015199134,"about_ca_system_score_gemma":0.0018424279,"threshold_uncertainty_score":0.021824539},"labels":[],"label_agreement":null},{"id":"W2049138229","doi":"10.1109/scam.2010.32","title":"Evaluating Code Clone Genealogies at Release Level: An Empirical Study","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Code refactoring; Software evolution; Software maintenance; Java; Computer science; Software system; Source code; Programming language; Software; Biology; Genetics; Software construction; Gene","score_opus":0.22123518092379388,"score_gpt":0.45418557628647194,"score_spread":0.23295039536267806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049138229","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99838233,0.00005501294,0.0010400801,0.000019905356,0.0000013697506,0.00002920314,0.00015733141,0.000021312826,0.00029340445],"genre_scores_gemma":[0.99709034,0.000044876673,0.0019221936,0.000012435755,0.000004020063,0.000041259067,0.0005839415,0.00001991831,0.00028104873],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9929542,0.0023864668,0.0006699516,0.0012665314,0.0024356626,0.00028712544],"domain_scores_gemma":[0.7049593,0.21034597,0.043403875,0.016654162,0.022211054,0.0024256753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009486735,0.00029162795,0.00035180885,0.0037736704,0.0007570924,0.0011998032,0.0008999633,0.0007372156,0.0008781696],"category_scores_gemma":[0.101905055,0.0003034498,0.00038609598,0.0027026194,0.0013600723,0.0026342159,0.0009658898,0.0011295951,0.00024394979],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020320485,0.00024981776,0.96606535,0.00006907752,0.00008708665,0.00027100326,0.0028051457,0.0018637605,0.0017876298,0.00032841967,0.00026463534,0.02600485],"study_design_scores_gemma":[0.000018694194,0.000630447,0.9805284,0.000030662497,0.00006448454,0.0006802146,0.0024539323,0.011460793,0.002649103,0.0003617769,0.0010865248,0.000035033503],"about_ca_topic_score_codex":0.003924025,"about_ca_topic_score_gemma":0.0058586164,"teacher_disagreement_score":0.009486735,"about_ca_system_score_codex":0.0008792619,"about_ca_system_score_gemma":0.0005479115,"threshold_uncertainty_score":0.050171256},"labels":[],"label_agreement":null},{"id":"W2049321492","doi":"10.1142/s0218194002000913","title":"ASSOCIATION ANALYSIS OF SOFTWARE MEASURES","year":2002,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software sizing; Software; Software metric; Software engineering; Software construction; Reusability; Regression testing; Software development; Software quality; Verification and validation; Association (psychology); Data mining; Statistics; Programming language; Mathematics","score_opus":0.014566029233017957,"score_gpt":0.24283468499269795,"score_spread":0.22826865575968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049321492","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27659467,0.0018630658,0.7106275,0.0005579483,0.0001469822,0.0004355105,0.002344166,0.00086271553,0.006567423],"genre_scores_gemma":[0.869709,0.00071251026,0.12412203,0.0001174636,0.0001794009,0.000785418,0.0030782425,0.000113353446,0.0011826173],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9737118,0.010476057,0.0024660514,0.005757618,0.0068759797,0.0007125657],"domain_scores_gemma":[0.8498331,0.10682705,0.017176611,0.013140628,0.011910715,0.0011119983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013472085,0.0009787901,0.0016223808,0.011866493,0.0010606884,0.0035402873,0.0010684166,0.00102941,0.0030949705],"category_scores_gemma":[0.09784903,0.00037988438,0.0015830178,0.011703285,0.0012665463,0.004698914,0.0025662566,0.0016199627,0.0009106081],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010735084,0.00039257883,0.4301268,0.0015585632,0.0026727382,0.0011427402,0.002433008,0.030136153,0.006590086,0.1093257,0.004214394,0.41033372],"study_design_scores_gemma":[0.000115580544,0.0012852775,0.25033936,0.00041833142,0.0021151416,0.0026369405,0.0025478574,0.29047012,0.01075734,0.41664115,0.022402028,0.00027083533],"about_ca_topic_score_codex":0.00056673225,"about_ca_topic_score_gemma":0.00038148463,"teacher_disagreement_score":0.013472085,"about_ca_system_score_codex":0.00083090167,"about_ca_system_score_gemma":0.00108634,"threshold_uncertainty_score":0.071247995},"labels":[],"label_agreement":null},{"id":"W2049463166","doi":"10.1145/774833.774840","title":"Plugging-in visualization","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Eclipse; Visualization; Computer science; Java; Plug-in; Software engineering; Process (computing); Software; World Wide Web; Programming language","score_opus":0.016047630399051736,"score_gpt":0.28819906185074123,"score_spread":0.2721514314516895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049463166","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015859304,0.0007494075,0.767993,0.0012068202,0.0008737355,0.00027207864,0.0022895592,0.18294828,0.027807862],"genre_scores_gemma":[0.21330355,0.001311758,0.7095756,0.0008824744,0.00030657972,0.00056469254,0.0076437406,0.035627257,0.030784424],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99848396,0.0004282877,0.000114878305,0.00024747496,0.00057482556,0.00015065027],"domain_scores_gemma":[0.9928455,0.0030591616,0.00027169465,0.0022292188,0.0012088516,0.00038545762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023639053,0.001926974,0.0008747906,0.002129125,0.000824415,0.004714958,0.002106572,0.0014503228,0.02159871],"category_scores_gemma":[0.013190372,0.0008621406,0.0012877409,0.0014422389,0.00056377344,0.004201849,0.0034823173,0.0020214028,0.0059127472],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009020861,0.00039660797,0.008179526,0.0013641329,0.00023826092,0.0027535446,0.0050700535,0.009169087,0.047227208,0.038253084,0.21651797,0.6699285],"study_design_scores_gemma":[0.00020602552,0.0002559134,0.0052807285,0.00044749718,0.0001872445,0.0022349874,0.00091667206,0.07044609,0.060830783,0.033429626,0.8254934,0.0002710423],"about_ca_topic_score_codex":0.0019968306,"about_ca_topic_score_gemma":0.0028469437,"teacher_disagreement_score":0.02159871,"about_ca_system_score_codex":0.00046098037,"about_ca_system_score_gemma":0.0008385905,"threshold_uncertainty_score":0.07225484},"labels":[],"label_agreement":null},{"id":"W2049511237","doi":"10.1016/j.scico.2012.11.002","title":"Stratified sampling of execution traces: Execution phases serving as strata","year":2012,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; Concordia University","funders":"","keywords":"Computer science; Sampling (signal processing); Stratified sampling; Parallel computing; Telecommunications; Statistics","score_opus":0.04349503827909484,"score_gpt":0.3280507423000704,"score_spread":0.2845557040209756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049511237","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.436699,0.00078148977,0.5481957,0.00065636553,0.00042981037,0.0035705399,0.0040591597,0.0014225073,0.004185344],"genre_scores_gemma":[0.9177194,0.00018452608,0.07281773,0.0006044356,0.00017868345,0.0018366325,0.0035854462,0.0001534373,0.0029198246],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97752863,0.013302934,0.0015795358,0.0035222587,0.0026582738,0.0014083599],"domain_scores_gemma":[0.9287406,0.03373099,0.0029944733,0.027256723,0.005427138,0.0018500523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026870728,0.00072480785,0.0018094184,0.001797788,0.0013617377,0.0026355453,0.0021940928,0.0015335038,0.0037836707],"category_scores_gemma":[0.10201899,0.0007480875,0.001589424,0.002049885,0.0016201417,0.0011429263,0.0021803635,0.0014393609,0.0010561921],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013590658,0.0020563758,0.55832356,0.0008487868,0.0024948325,0.000630351,0.007455528,0.02369402,0.030486159,0.08049674,0.01702056,0.26290238],"study_design_scores_gemma":[0.0026635053,0.0045339777,0.27907395,0.0003930205,0.0055501577,0.0013349869,0.004722344,0.3987767,0.051197503,0.20070964,0.050686568,0.00035760275],"about_ca_topic_score_codex":0.0043566595,"about_ca_topic_score_gemma":0.004561136,"teacher_disagreement_score":0.026870728,"about_ca_system_score_codex":0.0010191115,"about_ca_system_score_gemma":0.0026299136,"threshold_uncertainty_score":0.14210767},"labels":[],"label_agreement":null},{"id":"W2049800049","doi":"10.4236/ti.2013.44031","title":"Using the ISO 19761 COSMIC Measurement Standard to Reduce “Information Asymmetry” in Software Development Contracts and Enable Greater Competitiveness","year":2013,"lang":"en","type":"article","venue":"Technology and Investment","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Software; Information asymmetry; Software development; Quality (philosophy); Computer science; Lead (geology); Business; Industrial organization; COSMIC cancer database; Software quality; Environmental economics; Risk analysis (engineering); Economics; Finance; Operating system","score_opus":0.03153201895860408,"score_gpt":0.2521696490963532,"score_spread":0.22063763013774912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049800049","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11234289,0.002852522,0.5246314,0.003947873,0.0016103474,0.0012887418,0.0033784572,0.002747371,0.34720042],"genre_scores_gemma":[0.5444593,0.0027958658,0.41268337,0.0016344839,0.0005951399,0.0018116891,0.0058245813,0.0010415365,0.029154127],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9435366,0.0104020685,0.003935624,0.0018432928,0.03895166,0.0013307842],"domain_scores_gemma":[0.93360823,0.011984176,0.007743412,0.010374486,0.0355496,0.00074017496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0138812605,0.0011566415,0.0006282206,0.009481283,0.002003051,0.0049747108,0.0019283337,0.00216128,0.0033224626],"category_scores_gemma":[0.06189143,0.00047476962,0.0010080026,0.008892671,0.0029777635,0.0052686385,0.0021613992,0.001834777,0.0014080006],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002152649,0.000281045,0.03201474,0.00079018995,0.00006497046,0.00033345426,0.0032774804,0.010760248,0.017459137,0.5647382,0.040213037,0.3298521],"study_design_scores_gemma":[0.00009236637,0.00082780415,0.1197494,0.0017445516,0.0001920606,0.0016359235,0.0024931703,0.027541786,0.06434914,0.14776593,0.6327589,0.0008489674],"about_ca_topic_score_codex":0.019897854,"about_ca_topic_score_gemma":0.02241558,"teacher_disagreement_score":0.019897854,"about_ca_system_score_codex":0.006214521,"about_ca_system_score_gemma":0.009263631,"threshold_uncertainty_score":0.073412},"labels":[],"label_agreement":null},{"id":"W2049899452","doi":"10.1145/2652524.2652549","title":"Monitoring bottlenecks in achieving release readiness","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Portfolio; Context (archaeology); Product (mathematics); Software release life cycle; Revenue; Business; Computer science; Process (computing); Process management; Risk analysis (engineering); Software; Software quality; Software development; Accounting; Finance; Mathematics; Operating system","score_opus":0.014484797118494302,"score_gpt":0.26212688202245144,"score_spread":0.24764208490395714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049899452","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95400035,0.00096982915,0.03491245,0.0005387411,0.000064142376,0.0002444862,0.0011617634,0.0016179974,0.006490252],"genre_scores_gemma":[0.98665977,0.0002226075,0.011630216,0.00005391928,0.000020875026,0.000092700815,0.00073675625,0.00009946864,0.00048380916],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.994447,0.0012392206,0.0005978648,0.0012596913,0.0018940492,0.00056223845],"domain_scores_gemma":[0.9537086,0.014654739,0.016356941,0.0027498673,0.009583432,0.0029464008],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072862385,0.0009547071,0.00056897197,0.003612541,0.0007374873,0.002479698,0.0009467423,0.00065755274,0.0018158149],"category_scores_gemma":[0.04474976,0.0005443559,0.00029891275,0.0022946235,0.0004995876,0.0037638687,0.0016018801,0.0012960001,0.00095904083],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013225164,0.0006644983,0.7310853,0.00081865274,0.00020745266,0.0005288332,0.0036876162,0.02995802,0.016922494,0.0044249096,0.007322285,0.20305745],"study_design_scores_gemma":[0.000080038175,0.00181523,0.7498184,0.0004142916,0.00025441515,0.00068334356,0.004494066,0.20687333,0.018235723,0.006377982,0.010696165,0.00025710976],"about_ca_topic_score_codex":0.0062941713,"about_ca_topic_score_gemma":0.005865658,"teacher_disagreement_score":0.0072862385,"about_ca_system_score_codex":0.0010152203,"about_ca_system_score_gemma":0.0016804913,"threshold_uncertainty_score":0.038533747},"labels":[],"label_agreement":null},{"id":"W2050187629","doi":"10.1145/1052898.1052912","title":"Mylar","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":296,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"AspectJ; Computer science; Programmer; Java; Programming language; Task (project management); Context (archaeology); Code (set theory); Software engineering; Software system; Software; Aspect-oriented programming; Engineering","score_opus":0.013550619996946858,"score_gpt":0.25814201061575254,"score_spread":0.2445913906188057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050187629","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03487147,0.0027280557,0.20503245,0.012791337,0.003522645,0.00065472216,0.015019615,0.10086067,0.62451905],"genre_scores_gemma":[0.16942553,0.0023092204,0.16747727,0.007842579,0.0010209398,0.0010488561,0.018344034,0.017668044,0.6148636],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986369,0.00032549573,0.00010200165,0.00033329756,0.0004707431,0.00013157263],"domain_scores_gemma":[0.995129,0.0013999245,0.0003976133,0.0013033409,0.001344516,0.00042554963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001864478,0.00057695014,0.00044063717,0.0011492786,0.00088666455,0.0027377075,0.0015123752,0.0009795369,0.13628393],"category_scores_gemma":[0.011090889,0.00050335476,0.0005342382,0.0008208268,0.00064974185,0.004213873,0.0037121451,0.0012134117,0.07813773],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044739412,0.00009731516,0.0021841656,0.00046851186,0.000019528572,0.00023581172,0.0015343144,0.00040718005,0.0044908565,0.04014591,0.5716605,0.37830845],"study_design_scores_gemma":[0.000036642778,0.00010154989,0.0008687754,0.00013485661,0.000013581719,0.0003792506,0.00032586613,0.0012305088,0.0026257786,0.00597337,0.9882812,0.000028608081],"about_ca_topic_score_codex":0.00071663246,"about_ca_topic_score_gemma":0.0014746821,"teacher_disagreement_score":0.13628393,"about_ca_system_score_codex":0.000677774,"about_ca_system_score_gemma":0.00094468385,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2050349725","doi":"10.1145/775047.775095","title":"From run-time behavior to usage scenarios","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science","score_opus":0.024312565883526833,"score_gpt":0.2574888682438811,"score_spread":0.23317630236035428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050349725","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5171269,0.0010820625,0.46372408,0.0019603958,0.000039532897,0.00023320394,0.0018352552,0.005277695,0.008720946],"genre_scores_gemma":[0.9068456,0.0005519875,0.08922406,0.00012959584,0.000023799852,0.00015141511,0.001934206,0.0003137657,0.0008254379],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9959467,0.0014629354,0.0003879777,0.0007149892,0.0011912164,0.00029624358],"domain_scores_gemma":[0.97255075,0.018080048,0.0026428348,0.0041564447,0.0019960247,0.0005739822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035242892,0.0010577147,0.00048988446,0.0032539924,0.000632039,0.0036169,0.0014539161,0.0017834777,0.0013463182],"category_scores_gemma":[0.035887312,0.0008410264,0.0007940075,0.00202381,0.0016279079,0.005929189,0.0017846853,0.0017225024,0.00048764673],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012036875,0.00077122636,0.22112723,0.0011893533,0.00052959885,0.007766923,0.010595967,0.164733,0.02343334,0.09151973,0.006229475,0.47090048],"study_design_scores_gemma":[0.00006135056,0.00032567454,0.048949994,0.0003425797,0.0002122368,0.005715693,0.0032893626,0.7483433,0.017898807,0.15552953,0.019136362,0.00019507583],"about_ca_topic_score_codex":0.0034468703,"about_ca_topic_score_gemma":0.0031153373,"teacher_disagreement_score":0.0036169,"about_ca_system_score_codex":0.0009190663,"about_ca_system_score_gemma":0.0009944543,"threshold_uncertainty_score":0.018638432},"labels":[],"label_agreement":null},{"id":"W2050396504","doi":"","title":"Inferring semantically related words from software context","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; WordNet; Java; Code (set theory); Natural language processing; Software; Program comprehension; Context (archaeology); Information retrieval; Software maintenance; Artificial intelligence; Precision and recall; Programming language; Software development; Software system","score_opus":0.011341461763682397,"score_gpt":0.23698971394842105,"score_spread":0.22564825218473866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050396504","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.596864,0.0048203417,0.3673004,0.0006924039,0.00014662986,0.00041710155,0.005442858,0.014674815,0.00964141],"genre_scores_gemma":[0.8395816,0.0009300692,0.15144226,0.0001244541,0.000051961095,0.00010168057,0.005960512,0.00031294202,0.0014945555],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984836,0.00042865533,0.00016651272,0.00047975767,0.000336793,0.000104592866],"domain_scores_gemma":[0.9953904,0.0029029027,0.0004925136,0.00040867156,0.0007164759,0.000088930734],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00080676866,0.0012075469,0.0006159349,0.0070030233,0.0007680686,0.0013973102,0.0006923384,0.0010867516,0.004208192],"category_scores_gemma":[0.0098455185,0.00032626404,0.0009385197,0.0028799565,0.0005262869,0.0037222994,0.0019559276,0.00081368524,0.0022787114],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009433863,0.00034382576,0.06897208,0.002292646,0.00026906093,0.0009902022,0.0025398345,0.013578442,0.12458722,0.009184461,0.012029345,0.7642695],"study_design_scores_gemma":[0.00014277756,0.0006280188,0.062365815,0.0006123882,0.00067161984,0.005028047,0.005093479,0.6642023,0.1740721,0.042194083,0.044797108,0.00019225522],"about_ca_topic_score_codex":0.005077678,"about_ca_topic_score_gemma":0.0046135713,"teacher_disagreement_score":0.0070030233,"about_ca_system_score_codex":0.0005926394,"about_ca_system_score_gemma":0.0011872556,"threshold_uncertainty_score":0.014077783},"labels":[],"label_agreement":null},{"id":"W205097626","doi":"10.1007/978-3-319-13835-0_12","title":"Analysis and Improvement of Release Readiness – A Genetic Optimization Approach","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Identification (biology); Portfolio; Software deployment; Context (archaeology); Genetic algorithm; Product (mathematics); Process (computing); Interval (graph theory); Risk analysis (engineering); Operations research; Engineering; Machine learning; Mathematics; Business","score_opus":0.009495749141357611,"score_gpt":0.22542203081375534,"score_spread":0.21592628167239772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W205097626","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06939003,0.0006062865,0.92274284,0.00023764155,0.000046139663,0.00013964459,0.00006943288,0.00080982264,0.0059581944],"genre_scores_gemma":[0.659286,0.00049402606,0.33604723,0.00007603597,0.00003677231,0.0001732421,0.00016778009,0.00028203978,0.0034368655],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99907255,0.0002708373,0.000035340036,0.00017728827,0.00028614717,0.00015769518],"domain_scores_gemma":[0.997358,0.0017603186,0.0003426123,0.00013002242,0.00035029155,0.000058699785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021910155,0.0014543724,0.0013899243,0.0024382644,0.0005832143,0.001356148,0.0016857585,0.0013675081,0.0024838687],"category_scores_gemma":[0.005534834,0.00078027154,0.0018280265,0.0017086717,0.00080535456,0.0012506311,0.00068862055,0.0014045018,0.00041121608],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000057886024,0.00008407536,0.00071348046,0.00006858079,0.000046589215,0.000044127835,0.00004038645,0.9468331,0.0032028027,0.0041348925,0.00031601934,0.044458065],"study_design_scores_gemma":[0.0000071065415,0.000048472553,0.0003037004,0.000009399441,0.000029768189,0.000016090742,0.00001459559,0.99700934,0.0008794808,0.0014889406,0.000186274,0.0000068389627],"about_ca_topic_score_codex":0.0075778584,"about_ca_topic_score_gemma":0.0040044067,"teacher_disagreement_score":0.0075778584,"about_ca_system_score_codex":0.0013970374,"about_ca_system_score_gemma":0.0021259745,"threshold_uncertainty_score":0.015067518},"labels":[],"label_agreement":null},{"id":"W2051040119","doi":"10.1109/icpc.2010.25","title":"On the Comparability of Software Clustering Algorithms","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Comparability; Computer science; Data mining; Software; Decomposition; CURE data clustering algorithm; Algorithm; Canopy clustering algorithm; Correlation clustering; Machine learning; Mathematics; Programming language","score_opus":0.025698956950715238,"score_gpt":0.2766435082860594,"score_spread":0.25094455133534416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051040119","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49546733,0.008647936,0.46317223,0.0025961641,0.00074753875,0.0007893528,0.0017046069,0.0011550877,0.025719767],"genre_scores_gemma":[0.89243233,0.00080064987,0.10190142,0.0003826726,0.0004837513,0.00049002323,0.0024366311,0.00044920528,0.00062331965],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.80382466,0.11759292,0.014744585,0.01677546,0.044419684,0.0026426713],"domain_scores_gemma":[0.29282692,0.5951326,0.030976938,0.0478913,0.031170739,0.0020014802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.12372295,0.0015712768,0.0023001747,0.013513587,0.0025826355,0.009795362,0.0038241819,0.004390147,0.0022061644],"category_scores_gemma":[0.5090682,0.0006741539,0.0021895575,0.011070955,0.007374445,0.013480026,0.0057136957,0.0029610195,0.00067494204],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007583743,0.001029263,0.16775578,0.0032953087,0.00533952,0.0007548527,0.0056776865,0.19729687,0.0099409,0.16928586,0.0071864943,0.42485365],"study_design_scores_gemma":[0.0006996779,0.0045720385,0.08181152,0.0012418035,0.0013263841,0.0020050553,0.0032621082,0.41496953,0.022839807,0.45091045,0.015883826,0.00047788347],"about_ca_topic_score_codex":0.0010209638,"about_ca_topic_score_gemma":0.0007269695,"teacher_disagreement_score":0.12372295,"about_ca_system_score_codex":0.003391354,"about_ca_system_score_gemma":0.0017051056,"threshold_uncertainty_score":0.6543173},"labels":[],"label_agreement":null},{"id":"W2051162588","doi":"10.1007/s10664-014-9311-2","title":"Modelling the ‘hurried’ bug report reading process to summarize bug reports","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Process (computing); Sentence; Reading (process); Software bug; Quality (philosophy); Natural language processing; Artificial intelligence; Data science; Information retrieval; Software; Programming language; Linguistics","score_opus":0.02334213703925641,"score_gpt":0.2861944876746135,"score_spread":0.26285235063535706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051162588","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44097272,0.0007962775,0.54322314,0.0020539125,0.00016755122,0.00028242107,0.001369849,0.0028640945,0.008269981],"genre_scores_gemma":[0.93161625,0.00019946668,0.06218414,0.00008709277,0.000049775208,0.00010001551,0.00077458954,0.00018188288,0.004806778],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981724,0.00080582715,0.00014335568,0.0004371361,0.0002691523,0.00017205552],"domain_scores_gemma":[0.9665,0.025112558,0.003486669,0.0019020297,0.0023627107,0.0006360778],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042259595,0.0008133941,0.00071092503,0.0020083636,0.00040728628,0.0034535336,0.0014670703,0.0025634237,0.0060080932],"category_scores_gemma":[0.03912948,0.00080244103,0.00097105687,0.0016603683,0.00088999985,0.0028389222,0.001029312,0.0017833164,0.0014329753],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009620198,0.00047568124,0.039517857,0.00042191579,0.00021441848,0.0006760194,0.00263123,0.79119,0.0064581838,0.04873893,0.0042136293,0.10450011],"study_design_scores_gemma":[0.00002636349,0.00008667171,0.0032053678,0.000019217161,0.000038641072,0.000058612524,0.00008770356,0.9849858,0.0007634799,0.009844532,0.0008576593,0.00002601364],"about_ca_topic_score_codex":0.020050233,"about_ca_topic_score_gemma":0.014845828,"teacher_disagreement_score":0.020050233,"about_ca_system_score_codex":0.0012605392,"about_ca_system_score_gemma":0.0017350394,"threshold_uncertainty_score":0.039867043},"labels":[],"label_agreement":null},{"id":"W2051204868","doi":"10.1109/icsm.2012.6405249","title":"What makes a good code example?: A study of programming Q&amp;A in StackOverflow","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":290,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Documentation; Code (set theory); Answer set programming; Set (abstract data type); Questions and answers; World Wide Web; Programming language; Information retrieval","score_opus":0.07388383997217017,"score_gpt":0.3253173807742034,"score_spread":0.25143354080203323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051204868","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9902981,0.000120132354,0.0015126053,0.0023398886,0.00001678334,0.00008409239,0.000021928758,0.000016504995,0.0055900607],"genre_scores_gemma":[0.992916,0.00030234538,0.0026152513,0.0011212045,0.000014936687,0.00017735022,0.00003255358,0.000049143895,0.0027712681],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9702575,0.0243506,0.0006917838,0.0008888832,0.0024755427,0.0013356959],"domain_scores_gemma":[0.8380828,0.13979198,0.008349578,0.0022570395,0.0060263765,0.0054922504],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026797567,0.00045108612,0.00065525377,0.002117758,0.008639721,0.0066980803,0.002256479,0.0030294259,0.0036319073],"category_scores_gemma":[0.08899395,0.0010514105,0.00036678233,0.0017855936,0.0104644485,0.0099669,0.0054608667,0.0057832696,0.0004953208],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000029310037,0.00041836506,0.006434924,0.00013472822,0.000004508475,0.0005121239,0.98377717,0.000039777384,0.00038786946,0.0024605545,0.0007282776,0.005072325],"study_design_scores_gemma":[0.000024715042,0.00026265316,0.009087532,0.00028002993,0.000008585865,0.00038616612,0.9691765,0.00062689616,0.00044198186,0.0018638634,0.017805627,0.000035501995],"about_ca_topic_score_codex":0.00583446,"about_ca_topic_score_gemma":0.011364886,"teacher_disagreement_score":0.026797567,"about_ca_system_score_codex":0.0038875772,"about_ca_system_score_gemma":0.003741757,"threshold_uncertainty_score":0.14172077},"labels":[],"label_agreement":null},{"id":"W2051447211","doi":"10.1109/icpc.2010.13","title":"A Technique for Just-in-Time Clone Detection in Large Scale Systems","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Computer science; Software maintenance; Software; Software system; Operating system; Real-time computing; Biology","score_opus":0.00990881464862532,"score_gpt":0.26395759696610843,"score_spread":0.2540487823174831,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051447211","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009400943,0.00017701428,0.98182356,0.000097667675,0.00005208878,0.00013643935,0.000086395594,0.0074325213,0.00079347397],"genre_scores_gemma":[0.11199963,0.00014036828,0.8842894,0.00012705258,0.00006839523,0.00020945772,0.0002889521,0.0005899893,0.0022867098],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9955303,0.00062812184,0.00037662825,0.0008404045,0.002430444,0.00019397054],"domain_scores_gemma":[0.9806455,0.0064182826,0.0023232778,0.0071051433,0.0031471264,0.00036067583],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023431492,0.0009014532,0.0010867065,0.003482529,0.001354383,0.0016893848,0.0025964812,0.001686046,0.0025978005],"category_scores_gemma":[0.016814262,0.00085836434,0.00094042975,0.0043430645,0.0011640836,0.0050029303,0.002465886,0.0019006577,0.0012751877],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034828,0.00022504327,0.006216753,0.00036321787,0.00014090895,0.0005719404,0.0011939796,0.007833619,0.070377424,0.016786527,0.0073675886,0.88857466],"study_design_scores_gemma":[0.00037376705,0.0014697461,0.013597982,0.00024104498,0.0005375141,0.008659939,0.0006501807,0.55335134,0.25927448,0.07761301,0.0837778,0.0004531644],"about_ca_topic_score_codex":0.0016758252,"about_ca_topic_score_gemma":0.0021049788,"teacher_disagreement_score":0.003482529,"about_ca_system_score_codex":0.0006870879,"about_ca_system_score_gemma":0.0014306055,"threshold_uncertainty_score":0.012391925},"labels":[],"label_agreement":null},{"id":"W2051978688","doi":"10.1109/icsm.2010.5609530","title":"Revisiting common bug prediction findings using effort-aware models","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":208,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Predictive modelling; Eclipse; Software bug; Software quality; Process (computing); Software quality assurance; Quality (philosophy); Code (set theory); Software engineering; Machine learning; Data mining; Software; Software development; Programming language; Set (abstract data type)","score_opus":0.03014153758873253,"score_gpt":0.28279066796763735,"score_spread":0.2526491303789048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051978688","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6726855,0.005195181,0.2984278,0.009855234,0.00041120668,0.00022918342,0.0015386839,0.0025439195,0.009113293],"genre_scores_gemma":[0.9712503,0.00083154446,0.026183447,0.00032580478,0.00010743348,0.000051044703,0.00062007125,0.00018587323,0.00044445004],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98235947,0.0076722857,0.0012354275,0.00398667,0.0039454317,0.0008006455],"domain_scores_gemma":[0.6379564,0.30080047,0.019449675,0.021526955,0.018566081,0.0017004301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0319261,0.0024135755,0.0017435364,0.008125205,0.0011346639,0.0039787614,0.0033266589,0.0015725319,0.002941579],"category_scores_gemma":[0.18116371,0.00102283,0.0025083644,0.0066382606,0.0021748897,0.008968502,0.003019811,0.0038119364,0.00072217366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051210745,0.0006725083,0.5923176,0.0018721146,0.001673415,0.0008528302,0.003918084,0.11786027,0.0015783553,0.022046162,0.0077862525,0.24891041],"study_design_scores_gemma":[0.00010184,0.00038285524,0.119121954,0.0006792944,0.0007897105,0.0005066762,0.0015341786,0.82592356,0.0018706688,0.04481109,0.004112195,0.00016599013],"about_ca_topic_score_codex":0.027144251,"about_ca_topic_score_gemma":0.034871563,"teacher_disagreement_score":0.0319261,"about_ca_system_score_codex":0.002318522,"about_ca_system_score_gemma":0.0029163156,"threshold_uncertainty_score":0.16884333},"labels":[],"label_agreement":null},{"id":"W2052007825","doi":"10.1145/1287624.1287649","title":"Determining detailed structural correspondence for generalization tasks","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Generalization; Computer science; Redundancy (engineering); Task (project management); Artificial intelligence; Code (set theory); Machine learning; Programming language; Mathematics; Set (abstract data type); Engineering","score_opus":0.0249693261675011,"score_gpt":0.3101787220970322,"score_spread":0.2852093959295311,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2052007825","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06892218,0.00016262617,0.91114956,0.00026758265,0.000054218777,0.0007769374,0.0009435195,0.013119833,0.004603501],"genre_scores_gemma":[0.22174388,0.0001040883,0.7715443,0.00011278504,0.000027962695,0.0005037137,0.002625501,0.0019471307,0.0013906668],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9867531,0.0041963123,0.0015028905,0.0031738821,0.003813487,0.0005604289],"domain_scores_gemma":[0.9280522,0.040737655,0.007672053,0.013855134,0.008684816,0.0009980954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008953334,0.0020707943,0.0017172833,0.0041220575,0.0020273738,0.0027735217,0.0031689662,0.0027946436,0.010612052],"category_scores_gemma":[0.07733958,0.0011973326,0.0016146016,0.0019964809,0.0013640319,0.009161702,0.0051540504,0.002436924,0.004656122],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001045794,0.0006578474,0.0214264,0.0015719737,0.00015628307,0.0012014387,0.005292341,0.015823254,0.067863576,0.027122872,0.01841176,0.8394265],"study_design_scores_gemma":[0.00030050526,0.0013907433,0.037962962,0.00078792137,0.0003206109,0.0042317924,0.0071086194,0.55706877,0.17655827,0.1249796,0.08879874,0.0004914719],"about_ca_topic_score_codex":0.002119375,"about_ca_topic_score_gemma":0.003573153,"teacher_disagreement_score":0.010612052,"about_ca_system_score_codex":0.0011250178,"about_ca_system_score_gemma":0.003032642,"threshold_uncertainty_score":0.047350287},"labels":[],"label_agreement":null},{"id":"W2052382206","doi":"10.1109/cgames.2013.6632615","title":"Bringing auto dynamic difficulty to commercial games: A reusable design pattern based approach","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Sandbox (software development); Leverage (statistics); Software engineering; Limiting; Process (computing); Construct (python library); Software; Game design; Video game development; Quality (philosophy); Human–computer interaction; Programming language; Artificial intelligence; Engineering","score_opus":0.02443911031851611,"score_gpt":0.25315619533020434,"score_spread":0.22871708501168822,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2052382206","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007959732,0.000098148455,0.98685336,0.00044023973,0.00002556115,0.000308089,0.00005755429,0.0017647584,0.0024925966],"genre_scores_gemma":[0.047063265,0.00019941118,0.9485949,0.0001325922,0.00001618585,0.00038184496,0.00020489898,0.0006128929,0.00279394],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951357,0.001294252,0.0005771986,0.0008464684,0.0018260811,0.00032023498],"domain_scores_gemma":[0.9910477,0.0025631178,0.0010364425,0.003500583,0.001494551,0.00035767752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005685128,0.0015803581,0.0006276213,0.0031949452,0.00095067336,0.00367684,0.002887521,0.0017258617,0.0019453521],"category_scores_gemma":[0.013413052,0.0017511209,0.0023195404,0.0015259899,0.00246674,0.005375423,0.0037530686,0.0032082666,0.000949536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001852903,0.0009182771,0.012726721,0.0016730493,0.000385983,0.0017325748,0.009810118,0.032530643,0.06500934,0.23244789,0.009664601,0.63291556],"study_design_scores_gemma":[0.00023079236,0.0008500768,0.0060895057,0.00119671,0.0007423855,0.005817509,0.0032970046,0.3024208,0.077768065,0.23593675,0.36522967,0.00042075472],"about_ca_topic_score_codex":0.0033488174,"about_ca_topic_score_gemma":0.00589206,"teacher_disagreement_score":0.005685128,"about_ca_system_score_codex":0.0013816549,"about_ca_system_score_gemma":0.002661834,"threshold_uncertainty_score":0.030066133},"labels":[],"label_agreement":null},{"id":"W2052424802","doi":"10.1002/smr.274","title":"Feed‐forward and recurrent neural networks for source code informal information analysis","year":2003,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Source code; Recurrent neural network; Artificial neural network; Identifier; Sentence; Set (abstract data type); Domain (mathematical analysis); Connectionism; Context (archaeology); Process (computing); Content-addressable memory; Associative property; Generalization; Code (set theory); Preprocessor; Machine learning; Programming language","score_opus":0.026888661107951784,"score_gpt":0.31927112631382604,"score_spread":0.29238246520587424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2052424802","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12273186,0.00038269872,0.87118953,0.00022620695,0.000041923322,0.00009709138,0.0001684606,0.003688855,0.0014733421],"genre_scores_gemma":[0.80107766,0.00020338465,0.1956797,0.000055857592,0.00002501885,0.000118683005,0.0004867713,0.000112880654,0.0022400024],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99957925,0.00013956206,0.00003522579,0.00008755721,0.000110249486,0.000048104066],"domain_scores_gemma":[0.9984816,0.00085283775,0.0001858808,0.000111614936,0.0003379651,0.000030189634],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009840201,0.00072250684,0.00034159046,0.00090775423,0.00019330972,0.0005879825,0.00066026027,0.00052591355,0.0009805523],"category_scores_gemma":[0.0034251283,0.00034752235,0.0004820408,0.00058347057,0.00034658238,0.00094817765,0.0004613001,0.0008929546,0.0002410969],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021882029,0.00013957832,0.0025663215,0.00012557312,0.000110282104,0.00022538348,0.000202293,0.5632402,0.018695403,0.0032728945,0.0012326929,0.40997055],"study_design_scores_gemma":[0.0000016107736,0.00001063502,0.00025400185,0.0000029714035,0.0000048221996,0.0000039069128,0.0000067364244,0.9973206,0.0017190918,0.00055137364,0.00012092831,0.0000032798225],"about_ca_topic_score_codex":0.01117673,"about_ca_topic_score_gemma":0.011151084,"teacher_disagreement_score":0.01117673,"about_ca_system_score_codex":0.000843803,"about_ca_system_score_gemma":0.00068014656,"threshold_uncertainty_score":0.022223353},"labels":[],"label_agreement":null},{"id":"W2052468877","doi":"10.1109/icsme.2014.29","title":"CSCC: Simple, Efficient, Context Sensitive Code Completion","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Context (archaeology); Code (set theory); Set (abstract data type); Source code; State (computer science); Simple (philosophy); Programming language","score_opus":0.019581053871821488,"score_gpt":0.26506540034631465,"score_spread":0.24548434647449316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2052468877","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03955562,0.000587112,0.88446033,0.0003295071,0.00014136868,0.00073666364,0.00079960836,0.070022255,0.0033675914],"genre_scores_gemma":[0.17573634,0.00026194085,0.8113003,0.00021652564,0.0000740531,0.00032838,0.0026183808,0.0030132749,0.006450809],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946549,0.000840991,0.00024602166,0.00080147997,0.0032299273,0.00022665638],"domain_scores_gemma":[0.9865649,0.0033438453,0.001395696,0.0043728473,0.0036473249,0.00067539414],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025498196,0.0015603267,0.0010130797,0.0036071902,0.0009806145,0.0012889765,0.0028457674,0.0013438265,0.005499267],"category_scores_gemma":[0.022097312,0.00095397024,0.0010304722,0.0022118588,0.001231962,0.0029372133,0.003410775,0.0023591532,0.004440021],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004339266,0.0003893545,0.005297696,0.00033428104,0.000058212085,0.00021980936,0.00072924164,0.013610204,0.021103198,0.0050002504,0.031118719,0.9217051],"study_design_scores_gemma":[0.0002641431,0.0007116693,0.006323,0.00016635848,0.00009378115,0.0010024924,0.00054118177,0.8020312,0.08083576,0.016116915,0.09167601,0.0002375404],"about_ca_topic_score_codex":0.009891673,"about_ca_topic_score_gemma":0.011390239,"teacher_disagreement_score":0.009891673,"about_ca_system_score_codex":0.000874015,"about_ca_system_score_gemma":0.0038201516,"threshold_uncertainty_score":0.019668162},"labels":[],"label_agreement":null},{"id":"W2052776473","doi":"10.1007/s10664-014-9324-x","title":"A Large-Scale Empirical Study of the Relationship between Build Technology and Build Maintenance","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Software maintenance; Computer science; Deliverable; Abstraction; Software engineering; Software; Source code; Scale (ratio); Source lines of code; Code (set theory); Systems engineering; Software system; Engineering; Programming language","score_opus":0.026836341569620767,"score_gpt":0.2976999388033339,"score_spread":0.2708635972337131,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2052776473","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9969703,0.00021318528,0.00081529916,0.00013948762,0.0000048142574,0.000021180169,0.00015128753,0.000017463153,0.0016667956],"genre_scores_gemma":[0.99848807,0.00014433922,0.0006451793,0.000055731263,0.000009492316,0.00001710536,0.00024485908,0.0000067876836,0.00038839778],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9986411,0.0005935642,0.00008202243,0.00017356433,0.00041374261,0.000096060154],"domain_scores_gemma":[0.9019074,0.07541063,0.011575528,0.0043220217,0.0045682206,0.0022161158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002700774,0.00028766092,0.00024673907,0.0017194581,0.0008210525,0.0008258102,0.00065107807,0.0006782079,0.0028939424],"category_scores_gemma":[0.03367406,0.0003323713,0.00031780012,0.0022451538,0.00087354414,0.0013683151,0.00071592565,0.0012223044,0.000581597],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002459554,0.0023880396,0.9618664,0.00016655083,0.00024354868,0.00026019372,0.0017715917,0.0010079507,0.00095154863,0.0013941819,0.001428592,0.028275458],"study_design_scores_gemma":[0.000026393047,0.00041535654,0.99358714,0.000043480217,0.00009270181,0.0001787447,0.0015779784,0.0020669152,0.0004749079,0.00035198047,0.0011709661,0.000013450276],"about_ca_topic_score_codex":0.006204452,"about_ca_topic_score_gemma":0.013121573,"teacher_disagreement_score":0.006204452,"about_ca_system_score_codex":0.00069093774,"about_ca_system_score_gemma":0.0011041552,"threshold_uncertainty_score":0.01428324},"labels":[],"label_agreement":null},{"id":"W2053107307","doi":"10.1109/tse.2013.2297712","title":"Automatic Summarization of Bug Reports","year":2014,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":182,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of the Fraser Valley; University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Task (project management); Conversation; Software bug; Software; Quality (philosophy); Natural language processing; Information retrieval; Data science; World Wide Web; Software engineering; Programming language","score_opus":0.008384848850770735,"score_gpt":0.2187966960586759,"score_spread":0.21041184720790518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2053107307","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53976643,0.004430609,0.39539918,0.0013733484,0.00048392257,0.0013942878,0.010711393,0.040965296,0.005475547],"genre_scores_gemma":[0.6102859,0.0012796634,0.36072332,0.00018953404,0.000397109,0.0007776114,0.021217324,0.0013356162,0.003794006],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99554116,0.0020122086,0.00045516834,0.0007888114,0.0010482243,0.00015440061],"domain_scores_gemma":[0.9578745,0.023564234,0.0049713254,0.0034609116,0.009577771,0.0005512683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005584284,0.0013514715,0.0010474824,0.0054084975,0.00055270456,0.0016637373,0.0012127068,0.00080411305,0.0018433866],"category_scores_gemma":[0.044444524,0.0004912021,0.00057095475,0.001992469,0.00020929916,0.0018443248,0.0012024656,0.0008204041,0.0012433988],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010538487,0.00032694035,0.020277265,0.0024039813,0.00025207098,0.00062691147,0.0065053715,0.009192116,0.053792942,0.0016344304,0.026800338,0.87713385],"study_design_scores_gemma":[0.0005776219,0.0037143826,0.1258539,0.0011690594,0.0017175485,0.0028095755,0.0075268135,0.5076183,0.18475184,0.012223841,0.15150112,0.00053606863],"about_ca_topic_score_codex":0.001625836,"about_ca_topic_score_gemma":0.002365577,"teacher_disagreement_score":0.005584284,"about_ca_system_score_codex":0.00047597295,"about_ca_system_score_gemma":0.0010457858,"threshold_uncertainty_score":0.02953285},"labels":[],"label_agreement":null},{"id":"W2053112282","doi":"10.1109/scam.2014.15","title":"On the Use of Context in Recommending Exception Handling Code Examples","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code review; Exception handling; Code (set theory); Context (archaeology); Source code; Software; Static program analysis; Software quality; KPI-driven code analysis","score_opus":0.12807186808916818,"score_gpt":0.29384933654013995,"score_spread":0.16577746845097177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2053112282","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6378345,0.012757193,0.3353436,0.0008129155,0.000263814,0.0004937065,0.000842551,0.0029452618,0.008706443],"genre_scores_gemma":[0.8786351,0.0013033476,0.117723085,0.00020324907,0.00011720849,0.000089785215,0.0006674311,0.00008270321,0.0011779781],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984634,0.0004100396,0.00013706669,0.00049826124,0.00040298322,0.000088184344],"domain_scores_gemma":[0.99315965,0.0043043112,0.0005790438,0.0005329134,0.0012093206,0.00021475436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011937704,0.000982385,0.0008801818,0.005333174,0.00071981695,0.0011592988,0.0007424587,0.0009863432,0.000557728],"category_scores_gemma":[0.0091314865,0.00031237392,0.00061090343,0.0023813897,0.0004110932,0.0016510307,0.00068261864,0.0006892347,0.00035926627],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007549415,0.0006766274,0.124569684,0.00079701294,0.0005398208,0.00048507852,0.0009078247,0.034143317,0.015024389,0.0021762592,0.005432091,0.81449294],"study_design_scores_gemma":[0.00015486726,0.001004262,0.07603267,0.00036772006,0.0008295366,0.0011423269,0.0009866405,0.89094,0.010206059,0.0058786864,0.012258575,0.00019873214],"about_ca_topic_score_codex":0.010909564,"about_ca_topic_score_gemma":0.025454836,"teacher_disagreement_score":0.010909564,"about_ca_system_score_codex":0.00034040064,"about_ca_system_score_gemma":0.0008603451,"threshold_uncertainty_score":0.021692097},"labels":[],"label_agreement":null},{"id":"W2053264768","doi":"10.1145/1370256.1370279","title":"Towards a mutation-based automatic framework for evaluating code clone detection tools","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; clone (Java method); Cloning (programming); Code (set theory); Mutation; Frame (networking); Software engineering; Data mining; Machine learning; Programming language","score_opus":0.08639718895178901,"score_gpt":0.36100985446985684,"score_spread":0.27461266551806784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2053264768","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056261975,0.00043325944,0.9323552,0.00028374125,0.00004221767,0.0016748686,0.00030211126,0.0069020344,0.0017446458],"genre_scores_gemma":[0.1872796,0.000091340415,0.8103388,0.00007594777,0.000021757989,0.0010398843,0.00049449696,0.0003050467,0.00035318168],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94336903,0.022976296,0.005491712,0.0043268,0.022438869,0.0013974574],"domain_scores_gemma":[0.914919,0.037257113,0.01165049,0.008225782,0.0264735,0.0014740983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04580886,0.003335691,0.0038133,0.018161874,0.0021246152,0.00772154,0.005944145,0.0042253956,0.0015734192],"category_scores_gemma":[0.10115845,0.0011264639,0.002184523,0.005242081,0.003784039,0.0065423013,0.003919266,0.002324447,0.00057406456],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001698481,0.0035227379,0.046528663,0.001631609,0.000913652,0.0005049037,0.0013589766,0.22491223,0.12601222,0.055296578,0.0044221478,0.53319776],"study_design_scores_gemma":[0.00027218153,0.0017691599,0.008465567,0.00020069137,0.0001979196,0.00029720913,0.00039614187,0.9247723,0.04545091,0.015137168,0.0028071303,0.00023363746],"about_ca_topic_score_codex":0.0093708,"about_ca_topic_score_gemma":0.007007222,"teacher_disagreement_score":0.04580886,"about_ca_system_score_codex":0.005387038,"about_ca_system_score_gemma":0.006796189,"threshold_uncertainty_score":0.24226332},"labels":[],"label_agreement":null},{"id":"W2053287998","doi":"10.1109/cseet.2013.6595236","title":"Understanding individual contribution and collaboration in student software teams","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Team software process; Computer science; Process (computing); Visualization; Personal software process; TRACE (psycholinguistics); Software development; Collaborative software; Software; Work (physics); Knowledge management; Teamwork; Software development process; Software engineering; Engineering; Software construction","score_opus":0.033645179264747214,"score_gpt":0.2947431571742883,"score_spread":0.2610979779095411,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2053287998","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99082476,0.00008958058,0.0039855647,0.00036933337,0.0000095604655,0.000025639454,0.000007826172,0.000016930024,0.004670713],"genre_scores_gemma":[0.99848276,0.0000820081,0.0009110434,0.00003342731,0.000004294123,0.00002476736,0.000009157764,0.0000070144333,0.00044557103],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.98858184,0.008126315,0.00039618948,0.0008454896,0.0013816435,0.00066850695],"domain_scores_gemma":[0.9679192,0.02322039,0.003045097,0.0013765387,0.0021532287,0.0022856302],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011682407,0.0003890186,0.00046425287,0.002875172,0.0041101947,0.0073751155,0.0013863946,0.0012815911,0.0017387517],"category_scores_gemma":[0.03614003,0.0003569065,0.00031604405,0.0012253649,0.004245359,0.0059875627,0.007125405,0.0015366019,0.00029319836],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005766319,0.00022721136,0.071703464,0.00010179503,0.000029015817,0.00026139003,0.89639556,0.0005179324,0.0007797746,0.0049191993,0.00033452242,0.024672413],"study_design_scores_gemma":[0.000028000171,0.00021524844,0.04871238,0.00019980395,0.000047510177,0.00028270032,0.9132066,0.0054824613,0.0012241951,0.02190764,0.008652239,0.000041212734],"about_ca_topic_score_codex":0.0023120786,"about_ca_topic_score_gemma":0.0024644884,"teacher_disagreement_score":0.011682407,"about_ca_system_score_codex":0.0017969086,"about_ca_system_score_gemma":0.0023501576,"threshold_uncertainty_score":0.061783195},"labels":[],"label_agreement":null},{"id":"W2053465203","doi":"10.1145/1370114.1370130","title":"Promoting developer-specific awareness","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Context (archaeology); Computer science; Point (geometry); Space (punctuation); Code (set theory); Software; Source code; Software development; Software engineering; Computer security; World Wide Web; Knowledge management; Operating system; Programming language; Set (abstract data type)","score_opus":0.05321844019753847,"score_gpt":0.27176255480614997,"score_spread":0.2185441146086115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2053465203","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20368765,0.0008312715,0.68725884,0.008279762,0.00048328622,0.0009652362,0.00005513517,0.011396644,0.087042145],"genre_scores_gemma":[0.81752414,0.0005254525,0.16042855,0.0016774558,0.000210074,0.00039110723,0.00012453989,0.000506636,0.018611938],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9942913,0.0024510135,0.0003792656,0.0010783031,0.0012484575,0.00055160746],"domain_scores_gemma":[0.9608198,0.016599867,0.0045284345,0.008399289,0.005693873,0.0039588255],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009628035,0.0007435763,0.0004299665,0.0015580901,0.0013625751,0.0038889272,0.0012819433,0.0019648597,0.0036650246],"category_scores_gemma":[0.037898894,0.00076909474,0.00040728532,0.0006616265,0.001113882,0.007673852,0.007821015,0.002617284,0.0011767576],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003567156,0.0019321913,0.04921357,0.0010168029,0.00010951256,0.0011414149,0.051938213,0.0034928557,0.07765319,0.06959974,0.023766996,0.71977884],"study_design_scores_gemma":[0.0004269465,0.0024355568,0.064150125,0.0009949654,0.0005134529,0.0057572457,0.027107319,0.0564761,0.0727289,0.11249825,0.6563355,0.00057560904],"about_ca_topic_score_codex":0.00094003795,"about_ca_topic_score_gemma":0.0013412724,"teacher_disagreement_score":0.009628035,"about_ca_system_score_codex":0.00063364604,"about_ca_system_score_gemma":0.0032105958,"threshold_uncertainty_score":0.05091858},"labels":[],"label_agreement":null},{"id":"W2055083376","doi":"10.1109/iwsm.mensura.2014.31","title":"An Analogy-Based Approach to Estimation of Software Development Effort Using Categorical Data","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Analogy; Categorical variable; Fuzzy logic; Computer science; Data mining; Defuzzification; Cluster analysis; Fuzzy set; Fuzzy classification; Fuzzy set operations; Machine learning; Artificial intelligence; Algorithm; Fuzzy number","score_opus":0.06501548629685232,"score_gpt":0.3187146621619959,"score_spread":0.25369917586514357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2055083376","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07553058,0.00020400596,0.92062837,0.00016611733,0.000027991484,0.00020678963,0.00069355714,0.00071557437,0.001826986],"genre_scores_gemma":[0.51542574,0.00014059416,0.4818934,0.00004210763,0.000027047392,0.00041574755,0.0012461595,0.000046917463,0.0007623006],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948953,0.0023898738,0.0004301028,0.0008020525,0.0013656092,0.00011719767],"domain_scores_gemma":[0.98649466,0.008493938,0.0014896893,0.0015035926,0.0018726045,0.00014549276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004166377,0.00067789195,0.0007866241,0.0066627655,0.00048372455,0.001162354,0.0014300287,0.0008132581,0.0014237007],"category_scores_gemma":[0.028185438,0.0002448628,0.00080056133,0.005406959,0.0004582168,0.0019329744,0.0012575034,0.0011704608,0.00042692677],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037431438,0.0005573732,0.07210881,0.0006253486,0.0003073036,0.00020759914,0.0011604326,0.18251313,0.008023339,0.040345788,0.0034255458,0.690351],"study_design_scores_gemma":[0.00003444062,0.00022701756,0.020529369,0.00006043675,0.00004129738,0.00016389243,0.00029821837,0.9328885,0.0044844835,0.036712013,0.0044798655,0.00008054903],"about_ca_topic_score_codex":0.002220357,"about_ca_topic_score_gemma":0.0030410604,"teacher_disagreement_score":0.0066627655,"about_ca_system_score_codex":0.00075611495,"about_ca_system_score_gemma":0.0009465092,"threshold_uncertainty_score":0.022034168},"labels":[],"label_agreement":null},{"id":"W2055492252","doi":"10.1145/2076354.2076407","title":"SourceVis","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Visualization; Software visualization; Human–computer interaction; Software; Perspective (graphical); Software development; Software engineering; Data visualization; Software analytics; Software construction; Operating system; Data mining","score_opus":0.04424992678721978,"score_gpt":0.24383539782698918,"score_spread":0.1995854710397694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2055492252","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0048959795,0.0008105337,0.18223724,0.00096006424,0.00046546038,0.00050490466,0.036691923,0.6302838,0.1431501],"genre_scores_gemma":[0.06270525,0.001809208,0.25699317,0.0018323964,0.0003043534,0.0020468573,0.19007975,0.26475048,0.21947856],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986859,0.00017038027,0.00010238245,0.0002630125,0.0006599012,0.00011841129],"domain_scores_gemma":[0.9978994,0.00056883483,0.000090655834,0.0006159616,0.00066311343,0.00016210205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015275656,0.0012280839,0.0009569406,0.002219298,0.000893283,0.0035946425,0.0031699846,0.0015392159,0.10007502],"category_scores_gemma":[0.0059953434,0.0011027108,0.0010464939,0.001983136,0.00045664163,0.0039754235,0.0039167022,0.002261947,0.06290615],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007068588,0.00020091658,0.0013475958,0.0010118363,0.00016485387,0.00023343522,0.0008625345,0.00197731,0.007889511,0.02905493,0.7421738,0.21437633],"study_design_scores_gemma":[0.0001792607,0.000046662088,0.00077616144,0.00012367495,0.000032472475,0.00019396147,0.00009614388,0.008188964,0.0074879928,0.014335825,0.9684826,0.000056240046],"about_ca_topic_score_codex":0.0029694338,"about_ca_topic_score_gemma":0.003757182,"teacher_disagreement_score":0.10007502,"about_ca_system_score_codex":0.000729561,"about_ca_system_score_gemma":0.0014711238,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2055510056","doi":"10.1109/sera.2010.34","title":"An Approach for Detecting Execution Phases of a System for the Purpose of Program Comprehension","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"TRACE (psycholinguistics); Computer science; Program comprehension; Initialization; Reverse engineering; Software; Software system; Task (project management); Computation; Software development; Software engineering; Software maintenance; Software construction; Programming language; Theoretical computer science; Systems engineering","score_opus":0.03828904873783377,"score_gpt":0.32105251065612506,"score_spread":0.28276346191829127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2055510056","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005627197,0.0000635976,0.990581,0.00013754347,0.000011770367,0.00014267887,0.00007232275,0.0030647907,0.00029928298],"genre_scores_gemma":[0.035315447,0.00006514789,0.9634134,0.0000588725,0.0000145945805,0.00012596136,0.0002121925,0.00023162803,0.0005628089],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972857,0.0006272947,0.00023590433,0.0007662079,0.0008846602,0.00020024128],"domain_scores_gemma":[0.99130285,0.004007894,0.0013573758,0.0014820958,0.0015981476,0.00025160343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022178032,0.0022705866,0.001172015,0.005011906,0.001285065,0.0019319417,0.0021820145,0.0026526242,0.0028773411],"category_scores_gemma":[0.0123226745,0.0011065581,0.0016636044,0.002195208,0.0017644025,0.0040988177,0.0018017336,0.002931747,0.0012649298],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005800534,0.000552648,0.01054517,0.0009251132,0.00017128223,0.00065389957,0.0027484684,0.04376164,0.15061975,0.038165964,0.005868104,0.7454079],"study_design_scores_gemma":[0.00015821653,0.000666843,0.0042306655,0.00015295715,0.00022188609,0.0013926462,0.0005594049,0.7990118,0.12244931,0.04405126,0.026939295,0.00016574042],"about_ca_topic_score_codex":0.0043972144,"about_ca_topic_score_gemma":0.004980079,"teacher_disagreement_score":0.005011906,"about_ca_system_score_codex":0.00092894403,"about_ca_system_score_gemma":0.00346574,"threshold_uncertainty_score":0.011729002},"labels":[],"label_agreement":null},{"id":"W2055765785","doi":"10.1145/1774088.1774504","title":"Can complexity, coupling, and cohesion metrics be used as early indicators of vulnerabilities?","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":94,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Universität des Saarlandes; North Carolina State University","keywords":"Cohesion (chemistry); Secure coding; Computer science; Software security assurance; Software metric; Software; Software bug; Software development; Empirical research; Cyclomatic complexity; Security bug; Software quality; Software engineering; Computer security; Information security; Programming language; Statistics; Mathematics","score_opus":0.0328252097964023,"score_gpt":0.2922418905803097,"score_spread":0.2594166807839074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2055765785","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9253332,0.0019997125,0.062009376,0.0016588947,0.00013008557,0.00012526695,0.0008213636,0.00066866237,0.0072535807],"genre_scores_gemma":[0.9819546,0.00035203088,0.01667219,0.000080471735,0.000030096026,0.00007287501,0.00032253037,0.00007779605,0.0004374851],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99359965,0.0019857404,0.0005156154,0.00067748746,0.0028457507,0.0003758007],"domain_scores_gemma":[0.89910835,0.0458366,0.03436334,0.006922789,0.011749911,0.0020190547],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064877993,0.0014014562,0.0009661405,0.009997351,0.00046622855,0.0021314167,0.000782218,0.001207479,0.0007746148],"category_scores_gemma":[0.09534537,0.0005854186,0.0006365117,0.0067205345,0.0012469789,0.0060390695,0.0014414857,0.0012419393,0.0004818582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017151512,0.00015103453,0.8518898,0.00030476137,0.0004254747,0.00015378167,0.0014078093,0.009117121,0.0072967513,0.0034703752,0.0015457736,0.12406591],"study_design_scores_gemma":[0.000022820177,0.0006053594,0.93282896,0.00019798089,0.00015906329,0.00033588064,0.0010251541,0.041314244,0.0065647853,0.012945322,0.0038308674,0.00016960561],"about_ca_topic_score_codex":0.0048173317,"about_ca_topic_score_gemma":0.008094595,"teacher_disagreement_score":0.009997351,"about_ca_system_score_codex":0.0008656403,"about_ca_system_score_gemma":0.0011813286,"threshold_uncertainty_score":0.034311175},"labels":[],"label_agreement":null},{"id":"W2055777130","doi":"10.1007/s10664-015-9366-8","title":"Investigating technical and non-technical factors influencing modern code review","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":107,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Université de Montréal","funders":"","keywords":"Computer science; Code review; Process (computing); Code (set theory); Key (lock); Source code; Variety (cybernetics); Empirical research; Software engineering; Component (thermodynamics); Data science; Static program analysis; Software development; Computer security; Software; Artificial intelligence; Programming language","score_opus":0.05327895845453611,"score_gpt":0.31370679174594,"score_spread":0.2604278332914039,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2055777130","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98552126,0.0024286374,0.001735596,0.0019489169,0.00005005244,0.000085387896,0.00016868455,0.000049709888,0.008011801],"genre_scores_gemma":[0.9971733,0.0006378111,0.0008177828,0.00017349693,0.000044531145,0.000022472981,0.00010556763,0.000032555457,0.0009923936],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97902244,0.008113833,0.0022548165,0.0014885828,0.0076214746,0.0014988586],"domain_scores_gemma":[0.3626739,0.42215273,0.12708075,0.0099031385,0.069559604,0.008629865],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.022815673,0.00021242262,0.0003273773,0.0069289524,0.0014174397,0.004311479,0.0009769073,0.0008618725,0.0039581154],"category_scores_gemma":[0.35692203,0.00033275777,0.0004785622,0.006150205,0.0016077352,0.0036506841,0.0015031842,0.0015136322,0.0005186999],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003051897,0.00017139602,0.9261365,0.00047165676,0.00020282765,0.0002795087,0.005334309,0.0005329262,0.0012861066,0.002957347,0.001879202,0.060443003],"study_design_scores_gemma":[0.000015424559,0.00018728268,0.98319995,0.00026167452,0.00013048579,0.0003667104,0.005194132,0.0015114064,0.0012012728,0.0014032105,0.006487501,0.00004097562],"about_ca_topic_score_codex":0.010609989,"about_ca_topic_score_gemma":0.02444933,"teacher_disagreement_score":0.97718436,"about_ca_system_score_codex":0.0033962973,"about_ca_system_score_gemma":0.007930938,"threshold_uncertainty_score":0.12066227},"labels":[],"label_agreement":null},{"id":"W2055986279","doi":"10.1007/s11219-014-9233-7","title":"Prioritizing code-smells correction tasks using chemical reaction optimization","year":2014,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Code refactoring; Code smell; Computer science; Metaheuristic; Code (set theory); Search-based software engineering; Software; Software engineering; Software quality; Software design; Programming language; Artificial intelligence; Software development","score_opus":0.04201576491452648,"score_gpt":0.32948212129263316,"score_spread":0.2874663563781067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2055986279","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2984577,0.0009590948,0.67238605,0.0009029301,0.00022360579,0.00055096246,0.00033869888,0.017880494,0.008300344],"genre_scores_gemma":[0.6929396,0.0002675905,0.30071342,0.00021876946,0.00004113133,0.00012853205,0.0005079143,0.0009476161,0.0042354744],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987459,0.00025608373,0.00006394155,0.00024831758,0.00049705396,0.00018864665],"domain_scores_gemma":[0.99459594,0.0028722482,0.0006519823,0.00038824196,0.0012296156,0.00026196186],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011707608,0.0017215442,0.0012948822,0.0017491925,0.0006554086,0.001091673,0.001319401,0.00079605484,0.004201456],"category_scores_gemma":[0.0063091246,0.00049634784,0.0010013545,0.0010515422,0.0004282046,0.000886875,0.0009551197,0.0012277453,0.000756151],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013526626,0.0012682638,0.007773323,0.00078801933,0.00016013179,0.0002868393,0.00013253231,0.31089112,0.12829125,0.0029541936,0.005546606,0.540555],"study_design_scores_gemma":[0.000053147924,0.0003176613,0.0012466535,0.000013436439,0.000075924974,0.00004712347,0.00006427932,0.9481966,0.04677472,0.0017898458,0.0013975739,0.000023057371],"about_ca_topic_score_codex":0.0068268934,"about_ca_topic_score_gemma":0.012460351,"teacher_disagreement_score":0.0068268934,"about_ca_system_score_codex":0.00093014457,"about_ca_system_score_gemma":0.003988785,"threshold_uncertainty_score":0.014055252},"labels":[],"label_agreement":null},{"id":"W2056056849","doi":"10.1002/smr.277","title":"A user‐assisted approach to component clustering","year":2003,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Cluster analysis; Data mining; Component (thermodynamics); Partition (number theory); Graph; Software; Clustering coefficient; Visualization; Component-based software engineering; Theoretical computer science; Software system; Artificial intelligence; Programming language; Mathematics","score_opus":0.05873784470935019,"score_gpt":0.33637233252827137,"score_spread":0.2776344878189212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2056056849","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035525225,0.00003706909,0.99142843,0.00006767726,0.000013634134,0.00007067432,0.000052524556,0.0042677494,0.0005098162],"genre_scores_gemma":[0.071234144,0.000042693155,0.9262611,0.000057591555,0.000022895321,0.00017804361,0.00034083816,0.0005788543,0.0012837332],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9926569,0.003182771,0.00032422488,0.00096495845,0.0026944482,0.00017664825],"domain_scores_gemma":[0.98261756,0.006455434,0.00079301256,0.0051195105,0.0046837926,0.0003305877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052796234,0.0014368246,0.0012402756,0.00426477,0.0014255483,0.0027364462,0.0046769725,0.0018234045,0.005121848],"category_scores_gemma":[0.01941421,0.0008613343,0.001187747,0.0027108663,0.0011003164,0.0024784482,0.0030232393,0.0020718086,0.0023912454],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007295403,0.0005145402,0.0032626346,0.00033803197,0.00030357766,0.000335128,0.0018801909,0.08034322,0.030553719,0.027395824,0.016284443,0.8380591],"study_design_scores_gemma":[0.000051036437,0.00007466939,0.00060520833,0.000028482362,0.00003416332,0.00029418583,0.0001537626,0.95152336,0.017862605,0.014486252,0.014809249,0.00007697443],"about_ca_topic_score_codex":0.0026935223,"about_ca_topic_score_gemma":0.00471959,"teacher_disagreement_score":0.0052796234,"about_ca_system_score_codex":0.0008444593,"about_ca_system_score_gemma":0.0012355562,"threshold_uncertainty_score":0.027921677},"labels":[],"label_agreement":null},{"id":"W2056306711","doi":"10.1007/s11219-012-9180-0","title":"Influence of confirmation biases of developers on software quality: an empirical study","year":2012,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Assertion; Empirical research; Software quality; Confirmation bias; Software metric; Software development; Cognitive bias; Context (archaeology); Empirical evidence; Software bug; Software; Metric (unit); Quality (philosophy); Data science; Cognition; Software engineering; Psychology; Social psychology; Statistics; Engineering; Operations management; Mathematics; Programming language","score_opus":0.14038679657632722,"score_gpt":0.43106620055480366,"score_spread":0.29067940397847647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2056306711","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9980628,0.00020250598,0.0004278989,0.000110113724,0.000008624272,0.000018303346,0.000030059065,0.000012346502,0.0011274166],"genre_scores_gemma":[0.9994332,0.000055864944,0.00028981618,0.000029375291,0.000013840024,0.000008723443,0.000025371673,0.0000081857,0.00013572405],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.966521,0.019528765,0.002187338,0.0019552226,0.008443899,0.0013637702],"domain_scores_gemma":[0.18400733,0.7021711,0.06753559,0.01581276,0.025191758,0.005281414],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.029912265,0.00043352385,0.0005991871,0.0027812703,0.0010733722,0.0024013312,0.00093294354,0.0014922724,0.0023008196],"category_scores_gemma":[0.36258754,0.0003899301,0.00051750056,0.0017567246,0.0015492042,0.0019873732,0.0015390759,0.0015138282,0.0003405603],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011070311,0.0006016135,0.9707959,0.00014761266,0.00014698628,0.0003008158,0.004421147,0.00043081,0.0014686738,0.00042763434,0.0003713258,0.01978031],"study_design_scores_gemma":[0.00016658471,0.0011767464,0.9830056,0.00013428763,0.0004940955,0.00063135993,0.0046594357,0.0039187414,0.0034840615,0.0010508134,0.0012121798,0.00006604176],"about_ca_topic_score_codex":0.0041642347,"about_ca_topic_score_gemma":0.004319952,"teacher_disagreement_score":0.9700877,"about_ca_system_score_codex":0.0015610728,"about_ca_system_score_gemma":0.0031363023,"threshold_uncertainty_score":0.15819305},"labels":[],"label_agreement":null},{"id":"W2056570028","doi":"10.1109/esem.2011.24","title":"An Experimental Evaluation of the Impact of System Sequence Diagrams and System Operation Contracts on the Quality of the Domain Model","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Sequence diagram; Computer science; Unified Modeling Language; Software engineering; Software development process; Domain model; Software development; Domain (mathematical analysis); Software quality; Domain engineering; Software system; Process (computing); Use Case Diagram; Class diagram; Domain analysis; Goal-Driven Software Development Process; Software; Software construction; Programming language; Domain knowledge","score_opus":0.1732089469082812,"score_gpt":0.3896945407225418,"score_spread":0.21648559381426058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2056570028","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9926151,0.00013307811,0.0041219895,0.00008501928,0.00007294527,0.0007333228,0.00021474554,0.0001115029,0.0019122973],"genre_scores_gemma":[0.9780392,0.00024814985,0.01698813,0.000114245806,0.00006394519,0.0016263648,0.0005003683,0.00007589393,0.0023438062],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9921434,0.0035020078,0.00096288364,0.0014210446,0.0015413979,0.0004292982],"domain_scores_gemma":[0.80934465,0.16204318,0.0116511155,0.008681605,0.0056030788,0.0026764518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008548885,0.0014939358,0.0006645914,0.0008258011,0.00061077654,0.0011334448,0.0013907049,0.0019355568,0.006385288],"category_scores_gemma":[0.06194926,0.00073408545,0.00086790527,0.00072930416,0.0012424416,0.0019012666,0.0013401551,0.0016507846,0.00061220414],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.09983646,0.2881264,0.032596316,0.0060522063,0.0011341315,0.00086347095,0.011933008,0.032184906,0.31507608,0.0045108288,0.0034186034,0.20426767],"study_design_scores_gemma":[0.015675496,0.618032,0.09356558,0.0005632001,0.0016443525,0.0006200713,0.0034010352,0.062202763,0.1892149,0.0043037725,0.010380885,0.00039597307],"about_ca_topic_score_codex":0.0008594639,"about_ca_topic_score_gemma":0.000969551,"teacher_disagreement_score":0.008548885,"about_ca_system_score_codex":0.00077080313,"about_ca_system_score_gemma":0.0009481852,"threshold_uncertainty_score":0.045211375},"labels":[],"label_agreement":null},{"id":"W2056894403","doi":"10.1007/s10664-012-9231-y","title":"What are developers talking about? An analysis of topics and trends in Stack Overflow","year":2012,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":613,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Latent Dirichlet allocation; World Wide Web; Topic model; Popularity; Data science; Leverage (statistics); Android (operating system); Information retrieval; Artificial intelligence","score_opus":0.03051975966812,"score_gpt":0.3086728286717767,"score_spread":0.2781530690036567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2056894403","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.995082,0.00091489876,0.000580288,0.0008720514,0.00001657954,0.000017983984,0.00045520058,0.00003546379,0.0020255374],"genre_scores_gemma":[0.99575776,0.0012105553,0.0010224718,0.0001719724,0.00007407573,0.000033321685,0.00085620594,0.000052204992,0.00082148396],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99718946,0.0008733026,0.00034957347,0.0003361222,0.000933444,0.00031806846],"domain_scores_gemma":[0.9128747,0.056387976,0.01573795,0.0011747868,0.010728298,0.003096273],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0037436222,0.00022085135,0.00025373825,0.0086217085,0.0011083912,0.0021346668,0.00051845395,0.0007278773,0.0011728773],"category_scores_gemma":[0.041532904,0.0003245859,0.00029204867,0.008282863,0.0008052008,0.0043893224,0.001333264,0.0010974403,0.00027353564],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032378393,0.00013002686,0.8414324,0.00040087625,0.000055258413,0.00040967812,0.065325655,0.00022459858,0.0032429835,0.0014555064,0.0026769992,0.084322244],"study_design_scores_gemma":[0.000011042254,0.00010615956,0.9496954,0.00021500561,0.00007424177,0.00046980561,0.039001185,0.0013763714,0.0011681066,0.00070551404,0.0071460414,0.00003119017],"about_ca_topic_score_codex":0.009625689,"about_ca_topic_score_gemma":0.013259637,"teacher_disagreement_score":0.99625635,"about_ca_system_score_codex":0.0016067321,"about_ca_system_score_gemma":0.0024095215,"threshold_uncertainty_score":0.019798398},"labels":[],"label_agreement":null},{"id":"W2056908562","doi":"10.1109/idam.2014.6912671","title":"Preliminary study on design and development of a journal focused crawler system using EBD methodology: Part I &amp;#x2014; Design task and environment analysis","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"CERN","keywords":"Web crawler; Task (project management); Computer science; Domain (mathematical analysis); Identification (biology); Design cycle; Software engineering; Product design; Human–computer interaction; World Wide Web; Product (mathematics); Systems engineering; Engineering","score_opus":0.18903783844039185,"score_gpt":0.32197769944282806,"score_spread":0.1329398610024362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2056908562","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6040789,0.00037044,0.37026247,0.00052853936,0.00010025941,0.006320484,0.0007076787,0.0038130523,0.013818181],"genre_scores_gemma":[0.32310426,0.0003654672,0.6531934,0.000321137,0.000032852295,0.0030660296,0.0013416883,0.0008925306,0.017682627],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9947031,0.0026459026,0.0005194699,0.00075085653,0.0010817774,0.000298806],"domain_scores_gemma":[0.98563427,0.00748532,0.00065627054,0.0016497554,0.0037325372,0.0008419355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005789851,0.00064193807,0.00063263864,0.0016451613,0.0013222528,0.0016464321,0.0012034572,0.00097217795,0.0037400671],"category_scores_gemma":[0.012044947,0.0007558841,0.0006086436,0.0007265225,0.00063327537,0.0016392431,0.0009659325,0.00078648125,0.0012947669],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013695109,0.005313159,0.020749744,0.003302028,0.00016259521,0.0029068766,0.03707151,0.0134357,0.3118709,0.005538182,0.010792704,0.58748704],"study_design_scores_gemma":[0.0016001005,0.021669272,0.07899335,0.0010336337,0.0006193971,0.0040660044,0.023259733,0.12862466,0.43871978,0.0045165527,0.2964203,0.000477255],"about_ca_topic_score_codex":0.0022254623,"about_ca_topic_score_gemma":0.0036204415,"teacher_disagreement_score":0.005789851,"about_ca_system_score_codex":0.0011404788,"about_ca_system_score_gemma":0.002109745,"threshold_uncertainty_score":0.030619979},"labels":[],"label_agreement":null},{"id":"W2057128378","doi":"10.1109/iwsm-mensura.2013.35","title":"Measuring and Visualizing Code Stability -- A Case Study at Three Companies","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Volvo (Canada)","funders":"","keywords":"Visualization; Computer science; Product (mathematics); Quality (philosophy); Software; New product development; Quality assurance; Software engineering; Code (set theory); Stability (learning theory); Process management; Systems engineering; Engineering; Data mining; Business; Operations management; Operating system","score_opus":0.11679595665623213,"score_gpt":0.30954212441859946,"score_spread":0.19274616776236733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057128378","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9969874,0.000050989885,0.0022011253,0.00009822513,0.0000026249381,0.0000672677,0.00007459658,0.000062863066,0.00045501342],"genre_scores_gemma":[0.99133706,0.00008176155,0.007955818,0.00002703718,0.0000047617846,0.0000488588,0.000106783606,0.000035560417,0.00040236194],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9978403,0.0011163972,0.000117935226,0.0002344912,0.000500454,0.00019044847],"domain_scores_gemma":[0.9779093,0.015371977,0.0019352532,0.0013633448,0.0023709428,0.0010491998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003931934,0.0004368817,0.0002889922,0.002332882,0.0015414512,0.0013174013,0.00091975083,0.0013979338,0.0008094357],"category_scores_gemma":[0.01090188,0.00034502894,0.00038886952,0.0023389226,0.0013541257,0.0012125942,0.0010991137,0.0010616675,0.00018667644],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017734678,0.006678543,0.41362134,0.0011554349,0.00022434465,0.016468156,0.18652532,0.02412953,0.10093068,0.00283778,0.0037270691,0.24192834],"study_design_scores_gemma":[0.00027843245,0.0058418433,0.69826037,0.00028894717,0.0002313305,0.0047846637,0.11763284,0.06295036,0.09244029,0.0028129113,0.0141594615,0.00031853237],"about_ca_topic_score_codex":0.0085621225,"about_ca_topic_score_gemma":0.016432546,"teacher_disagreement_score":0.0085621225,"about_ca_system_score_codex":0.0010281747,"about_ca_system_score_gemma":0.00065404456,"threshold_uncertainty_score":0.020794332},"labels":[],"label_agreement":null},{"id":"W2057412652","doi":"10.1145/1188966.1188969","title":"Towards evidence-supported, question-directed collaborative program comprehension","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Program comprehension; Comprehension; Workflow; Documentation; Context (archaeology); Task (project management); Software engineering; Human–computer interaction; Process (computing); Knowledge management; Process management; Software; Software system; Systems engineering; Engineering; Programming language","score_opus":0.023290949404007277,"score_gpt":0.3168974156302802,"score_spread":0.2936064662262729,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057412652","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04310029,0.00026196678,0.9457239,0.0030726206,0.000028009337,0.0009706559,0.00012948744,0.0022545643,0.0044584638],"genre_scores_gemma":[0.17504582,0.00018490978,0.8227757,0.00025852266,0.000023301036,0.0006507477,0.00029356542,0.00010023195,0.0006671557],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9340204,0.04978809,0.0038779338,0.0049283896,0.0067190644,0.00066622277],"domain_scores_gemma":[0.6937256,0.24147679,0.016653636,0.022702908,0.023177542,0.002263552],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07657576,0.0017254214,0.001039414,0.0054322146,0.0016958306,0.011206247,0.0064078122,0.0053564273,0.0033682191],"category_scores_gemma":[0.25375953,0.0014832944,0.0012540974,0.002114537,0.00604539,0.014633254,0.008486932,0.0048011052,0.0014067026],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007166341,0.0026879048,0.02863314,0.004317721,0.0003769891,0.0016128436,0.13289069,0.019708982,0.026684556,0.091375776,0.0055501736,0.6854446],"study_design_scores_gemma":[0.0010199115,0.0021468094,0.015167296,0.0036637147,0.0006920099,0.0023275765,0.066642866,0.34001514,0.0824942,0.3875209,0.097719446,0.0005900583],"about_ca_topic_score_codex":0.0022431188,"about_ca_topic_score_gemma":0.0029339765,"teacher_disagreement_score":0.07657576,"about_ca_system_score_codex":0.0020557323,"about_ca_system_score_gemma":0.005688374,"threshold_uncertainty_score":0.4049762},"labels":[],"label_agreement":null},{"id":"W2057619340","doi":"10.1109/csmr-wcre.2014.6747182","title":"Examining the relationship between topic model similarity and software maintenance","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Source code; Software maintenance; Context (archaeology); Similarity (geometry); Metric (unit); Software; Software development; Software metric; Code (set theory); Software engineering; Relation (database); KPI-driven code analysis; Static program analysis; Code review; Software quality; Data mining; Programming language; Artificial intelligence; Engineering","score_opus":0.0911580655639271,"score_gpt":0.28506838849874844,"score_spread":0.19391032293482136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057619340","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9479552,0.0016108789,0.044741325,0.0010709031,0.000043868436,0.00009361536,0.00026867958,0.00016588981,0.004049612],"genre_scores_gemma":[0.99506265,0.00017542175,0.0041189985,0.000030119792,0.000035769914,0.000040626583,0.0003011892,0.000018936555,0.00021621471],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9924523,0.0047284164,0.00044521524,0.0010372351,0.00091993983,0.00041688708],"domain_scores_gemma":[0.6454638,0.3177373,0.020050261,0.0066885757,0.007067866,0.002992068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0139614,0.0006202334,0.00094973674,0.00484375,0.0012294948,0.004036891,0.0009837112,0.0018369954,0.00242471],"category_scores_gemma":[0.17857805,0.00042837646,0.0010552845,0.0068928264,0.0011320319,0.005197254,0.001974837,0.002298559,0.00040200324],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007437808,0.00043466725,0.92333996,0.00025075817,0.0008499527,0.00021771833,0.0035229796,0.013735266,0.0010698172,0.008330196,0.0012031216,0.046301693],"study_design_scores_gemma":[0.000072248185,0.0008625467,0.72268116,0.000099750614,0.0006322598,0.0009787441,0.0043081213,0.23530303,0.0009192434,0.03235149,0.001667482,0.00012405739],"about_ca_topic_score_codex":0.005581783,"about_ca_topic_score_gemma":0.004294408,"teacher_disagreement_score":0.0139614,"about_ca_system_score_codex":0.0014225689,"about_ca_system_score_gemma":0.0008239881,"threshold_uncertainty_score":0.07383585},"labels":[],"label_agreement":null},{"id":"W2057679337","doi":"10.1145/2593822.2593823","title":"Using developer conversations to resolve uncertainty in software development: a position paper","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Process (computing); Software; SPARK (programming language); Software development; Software development process; Position (finance); Analytics; Order (exchange); Software analytics; Data science; Software engineering; Knowledge management","score_opus":0.03099991737720705,"score_gpt":0.2789783397088573,"score_spread":0.24797842233165027,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057679337","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08577876,0.027804147,0.7744141,0.07356423,0.0028614681,0.000359658,0.0004377127,0.0011221668,0.033657745],"genre_scores_gemma":[0.637017,0.021029964,0.32332748,0.0035167036,0.0032718992,0.00041657835,0.0006490775,0.0006931305,0.010078193],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9791215,0.014896048,0.000660992,0.002345818,0.0024863197,0.000489262],"domain_scores_gemma":[0.8815003,0.101588584,0.0040178276,0.0039020914,0.0074091917,0.0015820442],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020235442,0.0012879604,0.0006359842,0.005396213,0.003328859,0.013333043,0.0020135008,0.0051604533,0.0033368405],"category_scores_gemma":[0.07642336,0.0012048371,0.0010999547,0.004276805,0.00505029,0.033399977,0.00529817,0.0046632495,0.0012051625],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005811119,0.00050583587,0.01972202,0.0014757504,0.00026601407,0.0005778154,0.05748575,0.013019777,0.003993504,0.30136696,0.026474727,0.5745308],"study_design_scores_gemma":[0.00014635066,0.00036987773,0.008952851,0.002228103,0.0002690987,0.0008300349,0.033423126,0.14947788,0.010186326,0.56185937,0.23169272,0.0005641984],"about_ca_topic_score_codex":0.0053116963,"about_ca_topic_score_gemma":0.0037018657,"teacher_disagreement_score":0.020235442,"about_ca_system_score_codex":0.0030777552,"about_ca_system_score_gemma":0.0034198866,"threshold_uncertainty_score":0.107016504},"labels":[],"label_agreement":null},{"id":"W2057780988","doi":"10.1016/j.jss.2007.05.035","title":"Applying machine learning to software fault-proneness prediction","year":2007,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":225,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"","keywords":"Software metric; Machine learning; Computer science; Support vector machine; Metric (unit); Software; Data mining; Artificial neural network; Artificial intelligence; Binary classification; Fault (geology); Software quality; Software fault tolerance; Reliability engineering; Software development; Engineering","score_opus":0.016601231710283004,"score_gpt":0.25773426312200903,"score_spread":0.24113303141172604,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057780988","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6125322,0.0013279619,0.38146457,0.0005967543,0.00019196335,0.00008348259,0.00024053866,0.0017188169,0.0018436373],"genre_scores_gemma":[0.96239454,0.00016197367,0.036575135,0.00005014089,0.000061460974,0.000020462225,0.00016974483,0.000026078602,0.0005403138],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990195,0.00041131355,0.0000982778,0.00014500962,0.0002426778,0.00008319113],"domain_scores_gemma":[0.98785573,0.009629927,0.0006756351,0.0004971676,0.0012091028,0.00013258537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018250488,0.00055242627,0.00080753514,0.001983202,0.00037731908,0.000691366,0.0007248268,0.0007645156,0.00067876565],"category_scores_gemma":[0.01340825,0.00024517937,0.00045895163,0.0011335188,0.00030066507,0.00080725626,0.0003754742,0.0007591676,0.00025156644],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026183954,0.00054476474,0.04681607,0.00013052911,0.00025283828,0.00013828799,0.000091970964,0.4723378,0.003745597,0.0007656368,0.0015925817,0.4733222],"study_design_scores_gemma":[0.0000076218885,0.000047523183,0.00261074,0.0000051655074,0.0000149747975,0.000023794802,0.000011442181,0.99387383,0.0013581481,0.0019297196,0.00011096696,0.00000612669],"about_ca_topic_score_codex":0.0066612125,"about_ca_topic_score_gemma":0.00498049,"teacher_disagreement_score":0.0066612125,"about_ca_system_score_codex":0.00048528533,"about_ca_system_score_gemma":0.00064188853,"threshold_uncertainty_score":0.013244867},"labels":[],"label_agreement":null},{"id":"W2058459251","doi":"10.1145/367845.368064","title":"A workbench for quality based software re-engineering (Doctoral Symposium)","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Workbench; Computer science; Software engineering; Process (computing); Quality (philosophy); Iterative and incremental development; Programming language; Object-oriented programming; Systems engineering; Engineering; Visualization; Artificial intelligence","score_opus":0.039297194479092835,"score_gpt":0.3021684074825268,"score_spread":0.26287121300343397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2058459251","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011689621,0.0013044975,0.93451273,0.0011795772,0.0008113725,0.0006175582,0.0005774128,0.018943256,0.030363966],"genre_scores_gemma":[0.051718503,0.0021216993,0.8416727,0.000393278,0.00035202855,0.00098399,0.0041504246,0.0061715418,0.09243573],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976305,0.000568461,0.00022075529,0.0003463588,0.0010579018,0.0001760886],"domain_scores_gemma":[0.99540114,0.0009287492,0.000200895,0.0016710123,0.0012501276,0.0005480483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005878485,0.0013262477,0.0008117943,0.0027013188,0.0013044669,0.00518659,0.0023736015,0.0015116603,0.023072356],"category_scores_gemma":[0.0063454392,0.0009137091,0.0009511237,0.0011814676,0.0011843916,0.0044927206,0.0039295666,0.002505307,0.013564517],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028019756,0.0005852112,0.0011578489,0.00050461374,0.00005729715,0.00053260056,0.0031447192,0.0063939653,0.07226869,0.07017454,0.06411473,0.7807857],"study_design_scores_gemma":[0.00015725636,0.00097134564,0.0017066787,0.0005523812,0.000075448974,0.0011855384,0.0009585979,0.027695052,0.081008464,0.042476635,0.84301984,0.00019269272],"about_ca_topic_score_codex":0.000939865,"about_ca_topic_score_gemma":0.000899329,"teacher_disagreement_score":0.023072356,"about_ca_system_score_codex":0.00097215164,"about_ca_system_score_gemma":0.0013493482,"threshold_uncertainty_score":0.07718474},"labels":[],"label_agreement":null},{"id":"W2059078574","doi":"10.1109/iecon.2012.6389404","title":"Case study: Using requirements and finite state machine for evaluating software trustworthiness","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Trustworthiness; Finite-state machine; Software engineering; Software; Software quality; Verification and validation; Software requirements specification; Software construction; Software verification; Software development; Programming language; Engineering; Computer security","score_opus":0.15096277830313692,"score_gpt":0.40139846410519736,"score_spread":0.25043568580206044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2059078574","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8728683,0.0001720525,0.11796616,0.00075497845,0.000051265724,0.0007559438,0.00029149652,0.00041067394,0.0067290976],"genre_scores_gemma":[0.9344166,0.0000842739,0.064195156,0.000052127376,0.000008837048,0.00026676786,0.00012083451,0.00003568725,0.00081977074],"study_design_codex":"simulation_or_modeling","study_design_gemma":"qualitative","domain_scores_codex":[0.9919735,0.0047064074,0.0004992441,0.00044795248,0.001996106,0.0003767412],"domain_scores_gemma":[0.9710697,0.022233335,0.0015342968,0.002204582,0.0022274063,0.0007306522],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005516194,0.00076277053,0.0005398924,0.0016855844,0.001120735,0.0012572891,0.0013374061,0.0026151985,0.0013138634],"category_scores_gemma":[0.021802895,0.00036286828,0.00078927306,0.0011088754,0.0013738327,0.0020921275,0.0013352475,0.0011948968,0.00027765374],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005043937,0.007894655,0.13676566,0.002957266,0.0008458738,0.05038094,0.02778332,0.33677888,0.10813503,0.10221036,0.00968908,0.21151502],"study_design_scores_gemma":[0.0007927954,0.007991803,0.03505549,0.00050797,0.0003892016,0.013724702,0.01247379,0.7509445,0.13132887,0.02641312,0.02003891,0.00033883352],"about_ca_topic_score_codex":0.0031641934,"about_ca_topic_score_gemma":0.0033028966,"teacher_disagreement_score":0.005516194,"about_ca_system_score_codex":0.0014234796,"about_ca_system_score_gemma":0.0011639134,"threshold_uncertainty_score":0.029172719},"labels":[],"label_agreement":null},{"id":"W2059600393","doi":"10.1145/2555596","title":"Predicting Stability of Open-Source Software Systems Using Combination of Bayesian Classifiers","year":2014,"lang":"en","type":"article","venue":"ACM Transactions on Management Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University; Université de Montréal","funders":"","keywords":"Interpretability; Machine learning; Computer science; Software quality; Artificial intelligence; Classifier (UML); Software system; Software; Software evolution; Data mining; Search-based software engineering; Stability (learning theory); Software sizing; Software metric; Component-based software engineering; Context (archaeology); Naive Bayes classifier; Bayesian probability; Software development; Software construction; Support vector machine; Operating system","score_opus":0.03265360708425621,"score_gpt":0.26780346713622266,"score_spread":0.23514986005196645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2059600393","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30247512,0.001138543,0.691578,0.00043134752,0.000089690264,0.0001918931,0.00035780706,0.0012223559,0.0025152038],"genre_scores_gemma":[0.88438386,0.0002983665,0.11337923,0.00008261865,0.00008938618,0.00009535421,0.00074517616,0.00005921208,0.00086684444],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99645776,0.0009401609,0.0003204348,0.0006854777,0.001314032,0.00028212383],"domain_scores_gemma":[0.9876001,0.007315063,0.0013845871,0.0006069697,0.002663428,0.0004297866],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005258295,0.0014411104,0.0014915038,0.0057812193,0.0008027423,0.0023204398,0.0012992233,0.0020763217,0.00077870686],"category_scores_gemma":[0.01724196,0.0006775244,0.0015384774,0.002059748,0.0004859637,0.0026989023,0.0013108312,0.001759643,0.00063590234],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005603141,0.00054739934,0.098005325,0.0001823667,0.0005549437,0.00022293736,0.0003732683,0.421396,0.006726931,0.002491849,0.0020658558,0.4668728],"study_design_scores_gemma":[0.000006115048,0.00006185039,0.0054127495,0.00001691007,0.000063621286,0.00003421643,0.000034917488,0.9903093,0.0013296253,0.0024022174,0.00031096608,0.000017616085],"about_ca_topic_score_codex":0.0065400084,"about_ca_topic_score_gemma":0.0069991844,"teacher_disagreement_score":0.0065400084,"about_ca_system_score_codex":0.0013339664,"about_ca_system_score_gemma":0.0010584439,"threshold_uncertainty_score":0.027808845},"labels":[],"label_agreement":null},{"id":"W2059657633","doi":"10.1109/wcre.2011.51","title":"Refactoring Traditional Forms into Ajax-enabled Forms","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Ajax; Code refactoring; Computer science; World Wide Web; Web application; Rich Internet application; Web page; Web development; Dynamic web page; The Internet; Web modeling; Programming language; Software","score_opus":0.06747865512665813,"score_gpt":0.2571120224965068,"score_spread":0.18963336736984865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2059657633","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0529922,0.00024355242,0.91072285,0.00039738085,0.0003293608,0.0007380816,0.00063664425,0.027131233,0.006808711],"genre_scores_gemma":[0.18112375,0.00051939284,0.7922777,0.00041061954,0.00011909271,0.0003438885,0.0020589477,0.005817171,0.017329438],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.997139,0.0004909727,0.0004055467,0.00046123017,0.0013412116,0.00016205192],"domain_scores_gemma":[0.98662657,0.0037765768,0.0011387325,0.0049838168,0.003216185,0.00025805188],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025804464,0.0012119169,0.00051542954,0.0012050173,0.0006122038,0.0017648654,0.0017797855,0.0009746978,0.0033887043],"category_scores_gemma":[0.017850729,0.00061696983,0.0009364773,0.0009336873,0.00097164646,0.0024169474,0.0020685152,0.0015705139,0.0020989967],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007479007,0.00070904486,0.012073078,0.0015238458,0.0001900312,0.0058167516,0.0045535173,0.028895102,0.12314334,0.04885267,0.028163942,0.74533075],"study_design_scores_gemma":[0.00037999294,0.0007811235,0.00582909,0.00047505868,0.00036181527,0.0050136354,0.00094250956,0.18165272,0.28238043,0.05111892,0.4707832,0.00028162336],"about_ca_topic_score_codex":0.0012808053,"about_ca_topic_score_gemma":0.0012924795,"teacher_disagreement_score":0.0033887043,"about_ca_system_score_codex":0.00047757218,"about_ca_system_score_gemma":0.0013562036,"threshold_uncertainty_score":0.013646901},"labels":[],"label_agreement":null},{"id":"W2059901918","doi":"10.1109/csmr-wcre.2014.6747161","title":"Automatic ranking of clones for refactoring through mining association rules","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; clone (Java method); Computer science; Similarity (geometry); Ranking (information retrieval); Data mining; Software; Programming language; Artificial intelligence; Biology; Genetics; Gene","score_opus":0.02692637272450593,"score_gpt":0.289120049818806,"score_spread":0.2621936770943001,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2059901918","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7273268,0.0025188494,0.25935346,0.00044127196,0.00008144523,0.0004916725,0.0026703067,0.004831857,0.0022844328],"genre_scores_gemma":[0.6028425,0.0006461079,0.38697225,0.000114031434,0.000042655895,0.00021824398,0.0076132254,0.00027381314,0.0012770878],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9936958,0.0010687311,0.0007552997,0.0014900716,0.0026302657,0.00035982236],"domain_scores_gemma":[0.96115017,0.021748286,0.0065713353,0.0030989116,0.006617297,0.00081400893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048015746,0.0015209124,0.0015774056,0.010206964,0.0009600295,0.002048243,0.0025276055,0.001308875,0.00080094446],"category_scores_gemma":[0.026028838,0.000578067,0.001652417,0.005343405,0.0004752342,0.0022440527,0.0010779538,0.0011082465,0.00067347445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047557065,0.0005766513,0.3420737,0.00079979823,0.00045660028,0.0016084168,0.0010785318,0.016015725,0.02386538,0.0017444419,0.003503583,0.6078015],"study_design_scores_gemma":[0.00017351154,0.0009392387,0.22586685,0.00054329896,0.0015825125,0.0054346295,0.001975557,0.6498046,0.08381739,0.008869857,0.020727566,0.00026503083],"about_ca_topic_score_codex":0.005354685,"about_ca_topic_score_gemma":0.00831583,"teacher_disagreement_score":0.010206964,"about_ca_system_score_codex":0.00072291214,"about_ca_system_score_gemma":0.0018012491,"threshold_uncertainty_score":0.025393426},"labels":[],"label_agreement":null},{"id":"W2060009252","doi":"10.1109/quatic.2010.11","title":"Experiments with Adding to the Experience that Can be Acquired from Software Courses","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Capstone; Software engineering; Capstone course; Course (navigation); Personal software process; Workflow; Context (archaeology); Software Engineering Process Group; Computer science; Social software engineering; Software peer review; Software development; Software construction; Engineering management; Software; Engineering; Programming language","score_opus":0.032460015642831075,"score_gpt":0.2949748955339821,"score_spread":0.262514879891151,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060009252","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9478031,0.00046129932,0.018749947,0.0007606657,0.00036314622,0.0014714822,0.00095619814,0.0011358269,0.028298361],"genre_scores_gemma":[0.89181876,0.00077367795,0.06918375,0.001132923,0.00016541671,0.002434067,0.0033306114,0.00056966895,0.030591024],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9946308,0.0024956877,0.00047202123,0.0008738296,0.0011091066,0.00041858983],"domain_scores_gemma":[0.95256615,0.033490244,0.0012818616,0.006862443,0.0028017757,0.0029976736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040371744,0.0011319488,0.0007394592,0.0009483132,0.00084628665,0.0024744892,0.002002264,0.002184328,0.017922698],"category_scores_gemma":[0.044901185,0.00047953601,0.00073871855,0.0008304746,0.00094164384,0.0036899385,0.0034026548,0.0022706424,0.004070143],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.016343305,0.06365907,0.031149713,0.0073422585,0.0003502752,0.0030325926,0.045033045,0.011980407,0.110629514,0.0091545535,0.03119447,0.67013085],"study_design_scores_gemma":[0.0050153485,0.099519804,0.11055581,0.0033095868,0.0016582777,0.006374813,0.045365598,0.061253987,0.18843679,0.029762369,0.4473938,0.0013538606],"about_ca_topic_score_codex":0.0006285042,"about_ca_topic_score_gemma":0.00084942323,"teacher_disagreement_score":0.017922698,"about_ca_system_score_codex":0.0005205691,"about_ca_system_score_gemma":0.0006540689,"threshold_uncertainty_score":0.059957385},"labels":[],"label_agreement":null},{"id":"W2060241165","doi":"10.1109/esem.2009.5316049","title":"Does explanation improve the acceptance of decision support for product release planning?","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Calgary","funders":"","keywords":"Product (mathematics); Computer science; Empirical research; Product planning; New product development; Decision support system; Knowledge management; Artificial intelligence; Marketing; Mathematics; Statistics; Business","score_opus":0.01592575360349527,"score_gpt":0.30149174307366455,"score_spread":0.2855659894701693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060241165","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97977495,0.00032468388,0.015260696,0.0007100816,0.000037662314,0.00058257236,0.00009346598,0.0006800447,0.0025358333],"genre_scores_gemma":[0.9593483,0.00028689724,0.03890054,0.00016886387,0.000034960936,0.00032669888,0.00017567707,0.0000497317,0.00070827303],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9862726,0.009830974,0.00074093667,0.0007983917,0.0019134483,0.00044358816],"domain_scores_gemma":[0.70490867,0.270062,0.014451083,0.0046305214,0.0047380575,0.0012096274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011340919,0.0009755563,0.0005703182,0.001176413,0.0003505761,0.0019526874,0.0011736131,0.0017633934,0.003008788],"category_scores_gemma":[0.12464861,0.00040670182,0.0006380983,0.0008201653,0.0006502966,0.0023249448,0.0011496596,0.0012692781,0.00035798622],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005147531,0.008569831,0.10624502,0.0047478387,0.00038537363,0.00066877174,0.0254866,0.0067839473,0.033249922,0.0011070836,0.0025856225,0.8050225],"study_design_scores_gemma":[0.0033380394,0.05147448,0.67896414,0.0026391463,0.0025666964,0.0017422811,0.030502623,0.12027207,0.078933135,0.005215254,0.023684205,0.0006678626],"about_ca_topic_score_codex":0.0010057214,"about_ca_topic_score_gemma":0.0013359179,"teacher_disagreement_score":0.011340919,"about_ca_system_score_codex":0.0007092526,"about_ca_system_score_gemma":0.0011113532,"threshold_uncertainty_score":0.059977233},"labels":[],"label_agreement":null},{"id":"W2060384068","doi":"10.1109/ictai.2011.155","title":"Machine-Learning Models for Software Quality: A Compromise between Performance and Intelligibility","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Compromise; Intelligibility (philosophy); Software quality; Machine learning; Software; Artificial intelligence; Domain engineering; Software engineering; Software construction; Software system; Software development","score_opus":0.13684350411366428,"score_gpt":0.3119661787820945,"score_spread":0.1751226746684302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060384068","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036747236,0.0015208771,0.9552407,0.0029851696,0.000043581294,0.00011410555,0.00010602195,0.00059236796,0.0026498334],"genre_scores_gemma":[0.7039416,0.0013613817,0.29226285,0.00039543686,0.00023282378,0.00035781343,0.00028113677,0.00015695431,0.0010100238],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9850604,0.00952624,0.0009475368,0.0010469394,0.0031707683,0.0002481254],"domain_scores_gemma":[0.8969095,0.09086204,0.0035974276,0.0043637566,0.0037478276,0.00051942054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024432145,0.0020402963,0.0020490286,0.003061786,0.00082037476,0.0069434466,0.0031081883,0.00350638,0.0011148475],"category_scores_gemma":[0.105964474,0.00079330185,0.0013975272,0.002068596,0.0033828008,0.009225216,0.0032145095,0.0045436956,0.0005544149],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003599463,0.00017096185,0.0073738834,0.00041868785,0.00018643732,0.00013530492,0.00089405285,0.7851533,0.0010989701,0.1011769,0.0011253206,0.10190622],"study_design_scores_gemma":[0.0000209085,0.00005683346,0.0005974753,0.00008531005,0.000028074457,0.000044037402,0.00005759878,0.89396423,0.00044885365,0.10413771,0.0005288269,0.000030072053],"about_ca_topic_score_codex":0.0013755156,"about_ca_topic_score_gemma":0.0012415218,"teacher_disagreement_score":0.024432145,"about_ca_system_score_codex":0.0018882997,"about_ca_system_score_gemma":0.0009585319,"threshold_uncertainty_score":0.12921101},"labels":[],"label_agreement":null},{"id":"W2060384944","doi":"10.1145/2597073.2597102","title":"Syntax errors just aren't natural: improving error reporting with language models","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Syntax error; Abstract syntax tree; Programming language; Syntax; Source code; Compiler; Java; Abstract syntax; Parsing; Consistency (knowledge bases); Exploit; Natural language processing; Artificial intelligence","score_opus":0.02666657826189673,"score_gpt":0.2828682662558349,"score_spread":0.25620168799393817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060384944","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033546068,0.00016740205,0.9471316,0.00078552193,0.000056365887,0.00016516531,0.00054396025,0.016747696,0.00085620355],"genre_scores_gemma":[0.267129,0.00023112638,0.7275997,0.00036862068,0.0000575212,0.0002674787,0.0013058012,0.0022224272,0.0008183762],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9840552,0.008602047,0.0012853335,0.0028246804,0.0028078537,0.00042479584],"domain_scores_gemma":[0.90155554,0.06695445,0.008828943,0.014279541,0.007608086,0.00077344093],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017575229,0.001873406,0.0013123646,0.0029912344,0.00097297534,0.004673979,0.003371477,0.002147369,0.0016067964],"category_scores_gemma":[0.10602789,0.0016692559,0.0020474603,0.002245454,0.0019096041,0.013963652,0.0040544053,0.0036834616,0.0014051386],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001113698,0.0011186749,0.047735874,0.0010718723,0.00049672625,0.0008896648,0.006126192,0.35997033,0.020375118,0.05405834,0.012683147,0.49436036],"study_design_scores_gemma":[0.000037073918,0.0000893417,0.0009395421,0.00007638777,0.0000616815,0.000114225215,0.0002384798,0.957206,0.0065287943,0.031638168,0.002999956,0.00007033944],"about_ca_topic_score_codex":0.0075515034,"about_ca_topic_score_gemma":0.0108982045,"teacher_disagreement_score":0.017575229,"about_ca_system_score_codex":0.0016882714,"about_ca_system_score_gemma":0.0039484566,"threshold_uncertainty_score":0.09294784},"labels":[],"label_agreement":null},{"id":"W2060590541","doi":"10.1109/csmr-wcre.2014.6747194","title":"Improving the detection accuracy of evolutionary coupling by measuring change correspondence","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Measure (data warehouse); Computer science; False positive paradox; Association (psychology); Coupling (piping); Association rule learning; Code (set theory); Subject (documents); Coupling strength; Artificial intelligence; Data mining; Programming language; Psychology; Engineering; Set (abstract data type)","score_opus":0.032688441919662393,"score_gpt":0.25000627491792565,"score_spread":0.21731783299826327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060590541","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5193718,0.0009492077,0.47099617,0.000332441,0.0001340095,0.00022976118,0.0005794289,0.004233466,0.0031736977],"genre_scores_gemma":[0.8436233,0.00015635647,0.1546832,0.00007547699,0.000047027836,0.00009413495,0.00060336664,0.00015786769,0.0005593056],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98444736,0.0029092657,0.0018291805,0.0032616286,0.006804892,0.0007476538],"domain_scores_gemma":[0.9162016,0.048117634,0.014258953,0.007746394,0.012063915,0.0016115106],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0079083275,0.0012223282,0.0016759058,0.012952317,0.0008833488,0.003036739,0.0025528313,0.0021917473,0.0012147119],"category_scores_gemma":[0.074844375,0.0006024994,0.0010207838,0.0063633774,0.0011609072,0.0049837735,0.0030973547,0.0015883181,0.000793221],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006825024,0.0005058937,0.56727463,0.0005753636,0.0005360005,0.00079853553,0.0012347595,0.033102665,0.024228446,0.0042010383,0.0023105224,0.36454967],"study_design_scores_gemma":[0.000072917275,0.0008235646,0.1922545,0.000092496964,0.00035668627,0.0027645412,0.00075892324,0.7340142,0.0567296,0.007711242,0.0042193555,0.00020200919],"about_ca_topic_score_codex":0.0026570205,"about_ca_topic_score_gemma":0.0022668906,"teacher_disagreement_score":0.012952317,"about_ca_system_score_codex":0.00074898096,"about_ca_system_score_gemma":0.0010556278,"threshold_uncertainty_score":0.041823745},"labels":[],"label_agreement":null},{"id":"W2060866317","doi":"10.1109/re.2012.6345841","title":"The quest for Ubiquity: A roadmap for software and systems traceability research","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Traceability; Requirements traceability; Computer science; TRACE (psycholinguistics); Field (mathematics); Change impact analysis; Software engineering; Data science; Software; Engineering management; Software development; Systems engineering; Engineering","score_opus":0.10959298506830575,"score_gpt":0.38726208526419437,"score_spread":0.27766910019588864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060866317","genre_codex":"commentary","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015605832,0.24919313,0.30528647,0.38891032,0.0031822524,0.00023121547,0.00041826328,0.0011361603,0.036036476],"genre_scores_gemma":[0.3319248,0.2819487,0.3378661,0.027319366,0.010435232,0.00069657335,0.0011462112,0.00094912114,0.007713847],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97388375,0.013346621,0.0016438825,0.003317845,0.0063964636,0.0014114485],"domain_scores_gemma":[0.7767348,0.17029676,0.006254359,0.022871682,0.016583951,0.0072584217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.045217857,0.0019307431,0.0031607132,0.011517515,0.004114892,0.01686499,0.0068123173,0.014612011,0.01706655],"category_scores_gemma":[0.07827453,0.0015950098,0.0023149576,0.010124565,0.022375813,0.097469606,0.015269687,0.014755293,0.0028338458],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016339094,0.00028516495,0.0029938756,0.0027694295,0.00008177213,0.00025801698,0.0018963671,0.0026829892,0.0010874612,0.67303336,0.015500569,0.2992476],"study_design_scores_gemma":[0.000028429982,0.00019967247,0.0013179986,0.0021253494,0.000047089165,0.00041783173,0.002890971,0.004474335,0.00056970096,0.84792614,0.13992448,0.00007803486],"about_ca_topic_score_codex":0.0034698066,"about_ca_topic_score_gemma":0.0020420211,"teacher_disagreement_score":0.045217857,"about_ca_system_score_codex":0.0058715534,"about_ca_system_score_gemma":0.0104809925,"threshold_uncertainty_score":0.23913777},"labels":[],"label_agreement":null},{"id":"W2060964996","doi":"10.1145/1321631.1321701","title":"Extracting rights and obligations from regulations","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Process (computing); Annotation; Government (linguistics); Stakeholder; Corporate governance; Business process reengineering; Focus (optics); Process management; Knowledge management; Business; Political science; Public relations; Artificial intelligence","score_opus":0.019290239150255656,"score_gpt":0.2826487268823837,"score_spread":0.263358487732128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060964996","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03046436,0.0003105815,0.9419226,0.0012693503,0.00017238951,0.0007999166,0.003873214,0.0043264595,0.016861219],"genre_scores_gemma":[0.14454675,0.00060413935,0.83609116,0.0003084971,0.00008194606,0.00083031075,0.011858383,0.0010040026,0.004674827],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9871596,0.0043979194,0.0013090235,0.0011582968,0.005506991,0.0004681516],"domain_scores_gemma":[0.9651091,0.01861329,0.002417585,0.005317115,0.008182525,0.00036041567],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007827735,0.0011311318,0.0008281241,0.006699708,0.0016518781,0.0037714164,0.0015505273,0.0018729613,0.003362783],"category_scores_gemma":[0.043555316,0.00096550013,0.0015768096,0.0026238894,0.0017168222,0.005705591,0.00427458,0.0022305558,0.002112805],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033541815,0.00042289653,0.010456304,0.0026612116,0.00013189907,0.0052634412,0.011827368,0.037073005,0.04075262,0.34698826,0.032666385,0.5114212],"study_design_scores_gemma":[0.00008601725,0.00012564073,0.004525336,0.0012116809,0.00017184725,0.001773202,0.0049524754,0.21681443,0.08428099,0.28792378,0.39788002,0.00025460235],"about_ca_topic_score_codex":0.005454531,"about_ca_topic_score_gemma":0.0061405827,"teacher_disagreement_score":0.007827735,"about_ca_system_score_codex":0.0015044092,"about_ca_system_score_gemma":0.00555762,"threshold_uncertainty_score":0.041397512},"labels":[],"label_agreement":null},{"id":"W2061072593","doi":"10.1007/s10664-013-9264-x","title":"SWordNet: Inferring semantically related words from software context","year":2013,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; WordNet; Program comprehension; Information retrieval; Software; Context (archaeology); Java; Ranking (information retrieval); Natural language processing; Code (set theory); Precision and recall; Software maintenance; Artificial intelligence; Software development; Programming language; Software system","score_opus":0.01656489003974,"score_gpt":0.2521907926293967,"score_spread":0.2356259025896567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061072593","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21641308,0.00352391,0.6703291,0.0012360181,0.0009333715,0.00080794404,0.02845489,0.06250844,0.015793275],"genre_scores_gemma":[0.4404377,0.0013915823,0.4990925,0.0005061411,0.0002479478,0.00059174455,0.05028054,0.00279609,0.004655763],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986125,0.00037252996,0.00015791129,0.0005079021,0.00025218853,0.000097077515],"domain_scores_gemma":[0.9970419,0.0018110266,0.00019191412,0.00049864687,0.00033806675,0.00011847414],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009817193,0.002267074,0.0011620562,0.0052760183,0.0014086638,0.0018138845,0.0015482788,0.001924787,0.008686532],"category_scores_gemma":[0.0063138558,0.0010282712,0.0015732356,0.0033536877,0.00092571386,0.008540773,0.0043499176,0.001687803,0.0039266692],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022355781,0.00085889647,0.022455452,0.0035993063,0.00076027703,0.0026863473,0.0033893643,0.015479437,0.051400114,0.05451141,0.08209397,0.76052994],"study_design_scores_gemma":[0.00048088137,0.0005886122,0.012556946,0.00061366806,0.0008967371,0.0019379924,0.004212559,0.5445394,0.039241817,0.28879505,0.105893634,0.00024266631],"about_ca_topic_score_codex":0.007772508,"about_ca_topic_score_gemma":0.014463442,"teacher_disagreement_score":0.008686532,"about_ca_system_score_codex":0.00078143796,"about_ca_system_score_gemma":0.0019265283,"threshold_uncertainty_score":0.02905935},"labels":[],"label_agreement":null},{"id":"W2061152427","doi":"10.1109/icpc.2013.6613831","title":"An empirical study on the efficiency of graphical vs. textual representations in requirements comprehension","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Comprehension; Representation (politics); Eye tracking; Process (computing); Graphical user interface; Contrast (vision); Software; Human–computer interaction; Visualization; Empirical research; Artificial intelligence; Natural language processing; Programming language; Mathematics; Statistics","score_opus":0.06396771391969484,"score_gpt":0.37427845089316064,"score_spread":0.3103107369734658,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061152427","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9969892,0.00013911696,0.0018895292,0.00003561786,0.0000030536387,0.000053963136,0.00005430723,0.0000240046,0.00081128796],"genre_scores_gemma":[0.99574834,0.00017141766,0.0032322572,0.000032529246,0.000010435485,0.00007670156,0.00017734276,0.000033573073,0.0005174075],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9910336,0.005545249,0.00095487386,0.0009903413,0.0012568601,0.00021902446],"domain_scores_gemma":[0.65892965,0.3113537,0.014460891,0.0057668975,0.008333969,0.001154818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011427773,0.00047313803,0.000364503,0.0011598018,0.00023722171,0.001165455,0.00051595963,0.00072350935,0.0024113022],"category_scores_gemma":[0.1057906,0.0002931732,0.0003181442,0.00086486404,0.0006760352,0.0017825306,0.00068343955,0.00057653757,0.00050982257],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043892316,0.0071422965,0.60490465,0.002172136,0.00043847566,0.00047240258,0.041467056,0.0026863962,0.05911814,0.0006324505,0.0014945348,0.27508226],"study_design_scores_gemma":[0.0003099365,0.00912309,0.9460152,0.00022481797,0.00032732106,0.00080550736,0.008632175,0.009217316,0.021176824,0.0006148758,0.0034551013,0.00009792763],"about_ca_topic_score_codex":0.00058974116,"about_ca_topic_score_gemma":0.00073373225,"teacher_disagreement_score":0.011427773,"about_ca_system_score_codex":0.00031608384,"about_ca_system_score_gemma":0.00031179254,"threshold_uncertainty_score":0.060436547},"labels":[],"label_agreement":null},{"id":"W2061250489","doi":"10.1109/icsm.2012.6405292","title":"Search-based refactoring: Towards semantics preservation","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Qatar Foundation; National Science Foundation","keywords":"Code refactoring; Computer science; Semantics (computer science); Programming language; Sorting; Domain (mathematical analysis); Program transformation; Software","score_opus":0.07054119937565884,"score_gpt":0.3131959618328126,"score_spread":0.24265476245715378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061250489","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09348663,0.00051656493,0.9018677,0.00038459408,0.00002657589,0.00023169485,0.00007984284,0.0010302819,0.0023760612],"genre_scores_gemma":[0.34338126,0.00037532096,0.6531201,0.00021212772,0.000019245343,0.0002330513,0.00027068163,0.0002912462,0.0020969335],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983808,0.00066510704,0.00008506406,0.00022944572,0.00052407425,0.000115547635],"domain_scores_gemma":[0.9979341,0.0009466102,0.00034017244,0.00029164305,0.0004077208,0.000079880934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029464036,0.0011314119,0.00093096297,0.0015124471,0.0004994503,0.00068506127,0.0014978296,0.0014127352,0.0010652259],"category_scores_gemma":[0.005949512,0.0005340865,0.0008687365,0.0011947817,0.00096321275,0.0011472179,0.0011549405,0.0010356895,0.0003646577],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018218657,0.00037252574,0.0034848466,0.00032041955,0.00014243911,0.00013430597,0.00037120504,0.67283744,0.029898841,0.008105549,0.0016415607,0.2825087],"study_design_scores_gemma":[0.00006815403,0.0001914501,0.00093180174,0.00004849567,0.000070084556,0.00011303542,0.0000781413,0.97676176,0.011269217,0.008032883,0.0024176985,0.000017175025],"about_ca_topic_score_codex":0.0041285376,"about_ca_topic_score_gemma":0.0056823143,"teacher_disagreement_score":0.0041285376,"about_ca_system_score_codex":0.00087188766,"about_ca_system_score_gemma":0.0023713964,"threshold_uncertainty_score":0.0155822635},"labels":[],"label_agreement":null},{"id":"W2061587515","doi":"10.1109/swan.2015.7070482","title":"Test case analytics: Mining test case traces to improve risk-driven testing","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Test (biology); Computer science; Analytics; Data science; Data mining; Geology","score_opus":0.059038662617794696,"score_gpt":0.3031889290144728,"score_spread":0.2441502663966781,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061587515","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.147921,0.001735063,0.8293845,0.00084456365,0.00010290913,0.0007594273,0.0023004632,0.014720346,0.002231782],"genre_scores_gemma":[0.5470201,0.00062594615,0.4440913,0.00019701933,0.00008851837,0.00051142904,0.0059496025,0.00058194116,0.0009341885],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960448,0.0009658927,0.00041282736,0.00060023996,0.0017837689,0.00019238284],"domain_scores_gemma":[0.9778691,0.013247826,0.0029084603,0.0018948549,0.0035494985,0.00053025014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038364467,0.0023828968,0.0012982857,0.010096428,0.00048717143,0.0018995048,0.0020780268,0.0009325306,0.0012116238],"category_scores_gemma":[0.02888908,0.0005258672,0.0012659897,0.004359057,0.0005857543,0.0028050467,0.0014421764,0.0013856865,0.0005617587],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005036928,0.0011139278,0.08785505,0.0009919973,0.0006545247,0.00077395653,0.0009094553,0.16772686,0.011180105,0.008270939,0.0091678025,0.7108517],"study_design_scores_gemma":[0.00005940173,0.00025565855,0.0075807804,0.00013786707,0.00009472114,0.00037750893,0.00019062615,0.9652251,0.007403061,0.015997345,0.002628992,0.000049003742],"about_ca_topic_score_codex":0.00672938,"about_ca_topic_score_gemma":0.007124619,"teacher_disagreement_score":0.010096428,"about_ca_system_score_codex":0.0010577242,"about_ca_system_score_gemma":0.0019313571,"threshold_uncertainty_score":0.020289361},"labels":[],"label_agreement":null},{"id":"W2061787477","doi":"10.1109/wcre.2009.59","title":"Ten Years Later, Experiments with Clustering as a Software Remodularization Method","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Cluster analysis; Computer science; Variety (cybernetics); Similarity (geometry); Data science; Software; Information retrieval; Artificial intelligence; Programming language","score_opus":0.017816400681915187,"score_gpt":0.304342125343366,"score_spread":0.2865257246614508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061787477","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8058176,0.0049828063,0.13485323,0.0024903866,0.0019173347,0.0028189893,0.004272833,0.008633304,0.034213576],"genre_scores_gemma":[0.693191,0.0015568997,0.275006,0.0010621003,0.00034180496,0.0019906461,0.0074135205,0.0014084773,0.018029489],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98490024,0.0077496064,0.0013612356,0.002246476,0.0031159394,0.0006265796],"domain_scores_gemma":[0.9129391,0.058297567,0.002224854,0.015614204,0.008843724,0.0020806629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011595116,0.0016269665,0.0014689818,0.0030382308,0.0019273455,0.0021651953,0.0019522709,0.0024248345,0.0049776784],"category_scores_gemma":[0.07747685,0.00061472936,0.0013561473,0.004527069,0.0012077758,0.005361793,0.00175888,0.0030383319,0.003355738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0054160478,0.013576965,0.015639773,0.0031137865,0.0010223622,0.0002502951,0.0036251415,0.043756694,0.023791933,0.009139677,0.051566515,0.82910085],"study_design_scores_gemma":[0.0025728995,0.026668383,0.07817302,0.0010915252,0.0015669619,0.0011808437,0.01031372,0.565499,0.12079106,0.04895995,0.14240544,0.00077714183],"about_ca_topic_score_codex":0.004686113,"about_ca_topic_score_gemma":0.007461102,"teacher_disagreement_score":0.011595116,"about_ca_system_score_codex":0.0017020312,"about_ca_system_score_gemma":0.0011326089,"threshold_uncertainty_score":0.061321557},"labels":[],"label_agreement":null},{"id":"W2061937509","doi":"10.1109/icsm.2011.6080782","title":"MoMS: Multi-objective miniaturization of software","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Miniaturization; Porting; Software; Software deployment; Process (computing); Software engineering; Embedded system; Operating system; Engineering; Electrical engineering","score_opus":0.036078041794740896,"score_gpt":0.2579216798817582,"score_spread":0.2218436380870173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061937509","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06647233,0.00020299364,0.9205562,0.0004473674,0.000024586665,0.0005133504,0.00008049911,0.0028417343,0.008860841],"genre_scores_gemma":[0.25550878,0.00018945466,0.7383346,0.00007919843,0.000013847858,0.0004517734,0.00018032794,0.00046815793,0.004773831],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99732494,0.0012069212,0.00015880309,0.00024888586,0.0008976077,0.00016276013],"domain_scores_gemma":[0.99568176,0.0024517067,0.0005654159,0.00080991536,0.0003782702,0.00011285306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002683515,0.0009652656,0.0003816021,0.00089351204,0.0005851667,0.0011973253,0.00092256075,0.0004398046,0.0043205917],"category_scores_gemma":[0.007722574,0.00053279445,0.00081592717,0.00040213086,0.0009680847,0.0015470476,0.002602294,0.00073934643,0.0005595435],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026294377,0.00039995534,0.0053850603,0.0011278802,0.00015483162,0.0008250961,0.004493201,0.0984006,0.123791166,0.07550859,0.0075224964,0.6821282],"study_design_scores_gemma":[0.00014890703,0.001156269,0.0066240514,0.00031087498,0.00018382566,0.0010456949,0.002260729,0.6828576,0.13768503,0.05163783,0.11596431,0.0001249294],"about_ca_topic_score_codex":0.0007546062,"about_ca_topic_score_gemma":0.0015337294,"teacher_disagreement_score":0.0043205917,"about_ca_system_score_codex":0.00071535277,"about_ca_system_score_gemma":0.00092620525,"threshold_uncertainty_score":0.014453828},"labels":[],"label_agreement":null},{"id":"W2062261433","doi":"10.1007/s00766-014-0205-z","title":"Patterns of continuous requirements clarification","year":2014,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Rework; Computer science; Requirements analysis; Requirements management; Oracle; Requirements engineering; User requirements document; Negotiation; Project management; Software engineering; Process management; Software; Data science; Systems engineering; Engineering","score_opus":0.025616603746807288,"score_gpt":0.2699037589057352,"score_spread":0.2442871551589279,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2062261433","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26841524,0.00031051555,0.6912845,0.0014970225,0.00012169546,0.0007157796,0.0006204361,0.0037160672,0.033318717],"genre_scores_gemma":[0.7485772,0.0001377934,0.23884098,0.00014875722,0.000026691667,0.00045913208,0.0011767951,0.0006127193,0.010020054],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9854431,0.005311594,0.001565346,0.0018606485,0.0050929626,0.00072639453],"domain_scores_gemma":[0.943977,0.028139846,0.005157641,0.01532333,0.006529066,0.00087316823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006267437,0.0005137371,0.00045292763,0.0017806989,0.0009275784,0.0033104287,0.0016546249,0.0016767324,0.005354083],"category_scores_gemma":[0.046465572,0.0008736673,0.0006853395,0.0020276748,0.0016241766,0.0057195304,0.0029021779,0.0018613675,0.0014309271],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008579832,0.0005153592,0.04380441,0.00082844775,0.00013106233,0.0024336826,0.02682905,0.012430192,0.038148273,0.28948116,0.011368231,0.5731722],"study_design_scores_gemma":[0.0002686656,0.0009995566,0.04049239,0.001090146,0.00025288045,0.0058022635,0.015245038,0.24985662,0.0590705,0.4419482,0.18464819,0.00032559363],"about_ca_topic_score_codex":0.0021060258,"about_ca_topic_score_gemma":0.002082405,"teacher_disagreement_score":0.006267437,"about_ca_system_score_codex":0.0009817858,"about_ca_system_score_gemma":0.0016163136,"threshold_uncertainty_score":0.033145785},"labels":[],"label_agreement":null},{"id":"W2062305127","doi":"10.1109/swan.2015.7070480","title":"Challenges and Issues of Mining Crash Reports","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Crash; Computer science; Data science; Programming language","score_opus":0.0889719980455779,"score_gpt":0.3225452044536358,"score_spread":0.23357320640805793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2062305127","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11994055,0.021409687,0.781087,0.055872146,0.00066851487,0.0013323145,0.008753992,0.0045350813,0.0064007347],"genre_scores_gemma":[0.2822924,0.008840739,0.696325,0.0016845736,0.0012851256,0.00087492977,0.0062879436,0.00060408533,0.0018051883],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.93219155,0.03338711,0.0105259,0.008116278,0.01463394,0.0011453029],"domain_scores_gemma":[0.58148044,0.3018099,0.031020783,0.036544725,0.04754117,0.0016029094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05460592,0.0014293069,0.002816632,0.018262176,0.0030393214,0.011701603,0.008755522,0.004152691,0.00087604637],"category_scores_gemma":[0.26355132,0.0019154951,0.0019217316,0.019836484,0.003917342,0.015867788,0.0048396117,0.0039488114,0.0012901175],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025755222,0.0003712772,0.13415498,0.003093164,0.0005087987,0.0013207849,0.008999325,0.02532673,0.0039956183,0.03379766,0.025217839,0.76295626],"study_design_scores_gemma":[0.00019884831,0.00048674372,0.12529793,0.0035366085,0.00068181776,0.0072247814,0.0314099,0.31374976,0.024476344,0.3271281,0.16504216,0.0007670198],"about_ca_topic_score_codex":0.008433373,"about_ca_topic_score_gemma":0.0067092795,"teacher_disagreement_score":0.05460592,"about_ca_system_score_codex":0.0015795511,"about_ca_system_score_gemma":0.004461,"threshold_uncertainty_score":0.28878713},"labels":[],"label_agreement":null},{"id":"W2062937278","doi":"10.1023/b:emse.0000048324.12188.a2","title":"An Empirical Exploration of the Distributions of the Chidamber and Kemerer Object-Oriented Metrics Suite","year":2004,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":71,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Collinearity; Suite; Computer science; Data mining; Correlation; Parametric statistics; Variance (accounting); Empirical research; Set (abstract data type); Test suite; Regression analysis; Statistics; Machine learning; Mathematics; Test case","score_opus":0.0281820532597251,"score_gpt":0.2982924489141738,"score_spread":0.27011039565444867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2062937278","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9164433,0.00083114253,0.07631865,0.0010597183,0.000016635322,0.000066115266,0.0005849218,0.00019548622,0.0044841375],"genre_scores_gemma":[0.98582876,0.0002103454,0.012537356,0.00005689905,0.00002014994,0.000047191174,0.00075115665,0.00008350517,0.00046473255],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98149085,0.011507259,0.0007536313,0.0015488712,0.004302619,0.0003967141],"domain_scores_gemma":[0.5578032,0.3902953,0.011434882,0.022157105,0.016227,0.002082508],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028988527,0.0005234363,0.0005536204,0.004895219,0.00070304354,0.0024080777,0.0019548424,0.0011193297,0.0022200085],"category_scores_gemma":[0.3402905,0.00040797252,0.00045686623,0.0049625062,0.002396953,0.0056342897,0.0020290646,0.002273678,0.00049777306],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000724419,0.00062859553,0.5852708,0.00025423968,0.00024809717,0.0004117664,0.006678092,0.046998706,0.0029213836,0.16160482,0.006714009,0.18754509],"study_design_scores_gemma":[0.00014483366,0.00074655045,0.36486873,0.0002013351,0.00011520736,0.0023522975,0.004154884,0.42747724,0.003851355,0.18393463,0.011999036,0.0001538804],"about_ca_topic_score_codex":0.0022539352,"about_ca_topic_score_gemma":0.002476842,"teacher_disagreement_score":0.97101146,"about_ca_system_score_codex":0.0015630599,"about_ca_system_score_gemma":0.0011017675,"threshold_uncertainty_score":0.1533078},"labels":[],"label_agreement":null},{"id":"W2063156085","doi":"10.1109/saner.2015.7081827","title":"Do code review practices impact design quality? A case study of the Qt, VTK, and ITK projects","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":136,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal","funders":"","keywords":"Code review; Software quality; Computer science; Software engineering; Code (set theory); Software; Quality (philosophy); Software peer review; Software construction; Software development; Set (abstract data type); Programming language","score_opus":0.36455942150848036,"score_gpt":0.4826728792808219,"score_spread":0.11811345777234156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2063156085","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9970192,0.00023946204,0.00090294064,0.0004974962,0.0000055089895,0.00008124605,0.000035238256,0.000012979661,0.001205975],"genre_scores_gemma":[0.9963881,0.0003393492,0.0023853611,0.00011473518,0.000008309217,0.000098580946,0.00005799919,0.000021550228,0.0005860122],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9692485,0.017324954,0.0019646867,0.0019892133,0.0077027553,0.0017700007],"domain_scores_gemma":[0.70792544,0.20355487,0.04735477,0.008304613,0.024967356,0.007893028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021783244,0.00040929875,0.0005009197,0.004682551,0.003505645,0.0029742138,0.0013647751,0.001600322,0.00086476764],"category_scores_gemma":[0.1265682,0.0007036954,0.0005255826,0.0047609964,0.0032940845,0.0036440196,0.003143553,0.0021811323,0.0002141667],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004883504,0.0019231226,0.6181662,0.00093462685,0.00016354525,0.0118584335,0.25881898,0.0018344158,0.0052408515,0.003202502,0.0025788173,0.09479019],"study_design_scores_gemma":[0.00013833265,0.0022426266,0.76455635,0.0007070918,0.00016582354,0.007071832,0.19975342,0.005535433,0.003972572,0.0017126333,0.013978359,0.00016551561],"about_ca_topic_score_codex":0.0155901015,"about_ca_topic_score_gemma":0.026849696,"teacher_disagreement_score":0.021783244,"about_ca_system_score_codex":0.005082571,"about_ca_system_score_gemma":0.0048535927,"threshold_uncertainty_score":0.11520219},"labels":[],"label_agreement":null},{"id":"W2063428900","doi":"10.1007/s10817-009-9121-1","title":"Interprocedural and Flow-Sensitive Type Analysis for Memory and Type Safety of C Code","year":2009,"lang":"en","type":"article","venue":"Journal of Automated Reasoning","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Memory safety; Type safety; Programming language; Type inference; Alias; Static analysis; Control flow; Source code; Data type; Compiler; Parallel computing; Algorithm; Inference; Artificial intelligence; Data mining","score_opus":0.010342703078796884,"score_gpt":0.28285217948543406,"score_spread":0.2725094764066372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2063428900","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14681265,0.00036230384,0.8451148,0.0003344124,0.000097213495,0.00011592246,0.00027765846,0.0037515687,0.0031334562],"genre_scores_gemma":[0.84465724,0.00013396183,0.15233538,0.0002204024,0.00009803203,0.000059650265,0.0002724507,0.00031046403,0.0019122767],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99570364,0.0007673799,0.00032799464,0.00053440913,0.0018772017,0.0007894595],"domain_scores_gemma":[0.98594695,0.006320867,0.0016062463,0.0036118627,0.0021888583,0.00032523912],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003739049,0.0008649358,0.00085120584,0.0041378764,0.0011916682,0.0020055901,0.0022556644,0.0012527156,0.002248979],"category_scores_gemma":[0.015274383,0.00079522963,0.0026920305,0.0014050665,0.0027023635,0.003171362,0.0021237694,0.0019279785,0.00036546472],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029792648,0.0005825654,0.042090528,0.0006138605,0.0005968431,0.0013154658,0.0012346996,0.18017909,0.06298578,0.22591816,0.0060564396,0.47544733],"study_design_scores_gemma":[0.00008523265,0.00025252157,0.004748845,0.00008111554,0.0003502292,0.00042662313,0.00016069457,0.7006409,0.07930924,0.21140781,0.0024200478,0.00011669626],"about_ca_topic_score_codex":0.0052554803,"about_ca_topic_score_gemma":0.004483838,"teacher_disagreement_score":0.0052554803,"about_ca_system_score_codex":0.0013488207,"about_ca_system_score_gemma":0.0032896884,"threshold_uncertainty_score":0.019774199},"labels":[],"label_agreement":null},{"id":"W2063462200","doi":"10.1007/s11219-012-9186-7","title":"Mining the impact of evolution categories on object-oriented metrics","year":2012,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cohesion (chemistry); Software evolution; Software; Computer science; Recall; Software metric; Software development; Data mining; Software engineering; Software construction; Programming language; Psychology; Cognitive psychology","score_opus":0.04837848210346001,"score_gpt":0.35755459381978144,"score_spread":0.3091761117163214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2063462200","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9885921,0.0009939127,0.0060300357,0.00022122521,0.000051599003,0.000050013994,0.0024461478,0.00024663316,0.0013682981],"genre_scores_gemma":[0.9837843,0.00032370508,0.008975265,0.00004280442,0.000045447894,0.000041284322,0.0061019133,0.00006591298,0.0006193097],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99437296,0.0012559787,0.0006737834,0.000795623,0.002546149,0.0003554268],"domain_scores_gemma":[0.961573,0.024614122,0.004624015,0.002111102,0.005821512,0.0012563185],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0032369094,0.00079440704,0.00075492996,0.014446359,0.0005727726,0.0015047437,0.0009240642,0.0008244275,0.000971784],"category_scores_gemma":[0.032217916,0.0002625728,0.001038044,0.009467813,0.0003550654,0.0021689464,0.0009711356,0.000905853,0.0004205107],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035997655,0.0004382124,0.8133389,0.00031757716,0.00044086116,0.00027485337,0.00034929757,0.0035074647,0.006186729,0.00087265874,0.0028283934,0.1710851],"study_design_scores_gemma":[0.00005577201,0.0007373326,0.88721716,0.00016143722,0.0006648742,0.00092056853,0.0009772258,0.09049374,0.0073513254,0.004889012,0.006477442,0.000054085005],"about_ca_topic_score_codex":0.005688388,"about_ca_topic_score_gemma":0.012971412,"teacher_disagreement_score":0.9967631,"about_ca_system_score_codex":0.0006774647,"about_ca_system_score_gemma":0.0011323489,"threshold_uncertainty_score":0.017118573},"labels":[],"label_agreement":null},{"id":"W2064015040","doi":"10.1145/859670.859695","title":"Extracting library-based Java applications","year":2003,"lang":"en","type":"article","venue":"Communications of the ACM","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Computer science; Java; Programming language; Operating system","score_opus":0.04705089450606053,"score_gpt":0.30345290938477926,"score_spread":0.2564020148787187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064015040","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054413706,0.0015420524,0.8431234,0.0007866578,0.00030448486,0.0005613809,0.003813344,0.088322125,0.0071328096],"genre_scores_gemma":[0.13224186,0.0013646793,0.82653004,0.0005217029,0.00018446498,0.00042197952,0.013501052,0.01181174,0.013422399],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998475,0.00017891367,0.00020685939,0.00029426566,0.0007071808,0.0001378066],"domain_scores_gemma":[0.99400604,0.0025640028,0.00045076766,0.0012968021,0.0015122783,0.00017012628],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008577282,0.0015392496,0.0010070149,0.0031992898,0.00069974316,0.0020222438,0.0016096369,0.0008074617,0.0044093365],"category_scores_gemma":[0.008796803,0.0009902706,0.0017331938,0.0029837026,0.0004595301,0.004354693,0.0028256476,0.0014428735,0.004488233],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045693462,0.00024922888,0.0060472176,0.0010452237,0.000115753945,0.00068298896,0.00039793306,0.0038848,0.12917197,0.0058084494,0.032982033,0.81915754],"study_design_scores_gemma":[0.00024226477,0.00036244153,0.012074675,0.0002623354,0.00042904567,0.0016766506,0.00051489484,0.13442558,0.5999492,0.027964417,0.22186802,0.0002304934],"about_ca_topic_score_codex":0.0015263392,"about_ca_topic_score_gemma":0.0032039555,"teacher_disagreement_score":0.0044093365,"about_ca_system_score_codex":0.0004385894,"about_ca_system_score_gemma":0.0016549645,"threshold_uncertainty_score":0.014750719},"labels":[],"label_agreement":null},{"id":"W2064132582","doi":"10.1007/s11219-010-9124-5","title":"Identification and analysis of attributes and base measures within ISO 9126","year":2010,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Measure (data warehouse); Software; Identification (biology); Task (project management); Software quality; Quality (philosophy); Reliability engineering; Set (abstract data type); Base (topology); Engineering; Computer science; Process (computing); Data mining; Systems engineering; Mathematics; Software development; Programming language","score_opus":0.03959637261930754,"score_gpt":0.3236045379643921,"score_spread":0.28400816534508455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064132582","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8245511,0.00092845043,0.16010831,0.0002534023,0.000078761026,0.0004937911,0.0021121886,0.0008387227,0.010635205],"genre_scores_gemma":[0.9062373,0.00022541113,0.0890246,0.00002525531,0.000017295211,0.000307313,0.003114473,0.000119169665,0.0009291978],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9879502,0.0016104956,0.0012948596,0.0006550555,0.0080818245,0.00040758343],"domain_scores_gemma":[0.959201,0.014012897,0.004352208,0.004103525,0.017680487,0.0006498202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008421085,0.000489884,0.0007519263,0.009011422,0.00093374564,0.0030801878,0.0010218546,0.0005603817,0.00066501077],"category_scores_gemma":[0.036267344,0.00032583726,0.00074868585,0.008863719,0.00062796264,0.0023417461,0.0011098371,0.0010864856,0.0003993863],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009644648,0.0009694057,0.32551637,0.0008723681,0.00017549992,0.000362075,0.0033037283,0.014626837,0.032828357,0.030758012,0.0030344801,0.58658844],"study_design_scores_gemma":[0.000094825635,0.0023639651,0.68604505,0.0005607279,0.00061217346,0.000982059,0.004639295,0.16907193,0.070124,0.030349385,0.03494305,0.00021358836],"about_ca_topic_score_codex":0.0037436294,"about_ca_topic_score_gemma":0.002972967,"teacher_disagreement_score":0.009011422,"about_ca_system_score_codex":0.0016759863,"about_ca_system_score_gemma":0.0029284717,"threshold_uncertainty_score":0.044535458},"labels":[],"label_agreement":null},{"id":"W2064432516","doi":"10.1109/saner.2015.7081891","title":"Towards a framework for automatic correction of anti-patterns","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code refactoring; Heuristics; Computer science; Task (project management); Software engineering; Software maintenance; Quality (philosophy); Test suite; Suite; Software evolution; Software quality; Software; Software design pattern; Software development; Test case; Machine learning; Software construction; Programming language; Systems engineering; Engineering","score_opus":0.04094471871904688,"score_gpt":0.3203529518301031,"score_spread":0.2794082331110562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064432516","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008100247,0.00013868835,0.99506134,0.00016851246,0.000021386884,0.000114839044,0.000040828,0.003240117,0.0004042213],"genre_scores_gemma":[0.013734904,0.00013331044,0.98475915,0.000098282246,0.000028834696,0.00014664611,0.00018564805,0.00036836535,0.0005448195],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98875254,0.00269749,0.0011709118,0.0023992546,0.004226811,0.00075304904],"domain_scores_gemma":[0.9805808,0.007384509,0.0022712431,0.0048180353,0.0043110936,0.00063434604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011433269,0.0021378754,0.0020925086,0.0061079105,0.0015665098,0.005218177,0.007625886,0.004308281,0.0031836927],"category_scores_gemma":[0.02995591,0.0024445574,0.0045828894,0.003353772,0.0042162407,0.005649966,0.004908717,0.0051161754,0.0022528716],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031497938,0.00056241825,0.0060938764,0.0014909131,0.0004399358,0.001234512,0.0015885656,0.1858608,0.026030786,0.25722566,0.012445481,0.50671214],"study_design_scores_gemma":[0.000090139416,0.000161648,0.0005839812,0.00036363656,0.00016223008,0.00055723137,0.00016625284,0.78448606,0.012272787,0.1674543,0.033595458,0.00010631335],"about_ca_topic_score_codex":0.007757155,"about_ca_topic_score_gemma":0.008770221,"teacher_disagreement_score":0.011433269,"about_ca_system_score_codex":0.002033034,"about_ca_system_score_gemma":0.0056784255,"threshold_uncertainty_score":0.060465634},"labels":[],"label_agreement":null},{"id":"W2064514756","doi":"10.5555/2819303.2819311","title":"Planning for the unknown: lessons learned from ten months of non-participant exploratory observations in the industry","year":2015,"lang":"en","type":"article","venue":"Conducting Empirical Studies in Industry","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Exploratory research; Observational study; Process (computing); Knowledge management; Computer science; Process management; Management science; Engineering; Medicine; Sociology","score_opus":0.8196605695777701,"score_gpt":0.5127729672968058,"score_spread":0.30688760228096434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064514756","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8704952,0.0031228273,0.081971616,0.022525348,0.0007350159,0.0031447606,0.0005109949,0.0006490849,0.016845217],"genre_scores_gemma":[0.9533647,0.0016905599,0.03666483,0.0023836861,0.00016908486,0.001864262,0.00027759746,0.0002527004,0.0033325248],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9310394,0.056265764,0.001816086,0.003288982,0.0045956355,0.002994144],"domain_scores_gemma":[0.7027379,0.22938818,0.00850025,0.020095699,0.028371803,0.010906203],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09336447,0.0013332239,0.0016196674,0.0017289075,0.0103726005,0.0071842954,0.0061560166,0.003761064,0.0032705863],"category_scores_gemma":[0.18765654,0.0012017454,0.00088335766,0.0016927143,0.009117513,0.013287446,0.008950082,0.0069730408,0.0013799295],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046007748,0.0019723715,0.025573926,0.0011791807,0.00007372178,0.0028351326,0.8300124,0.0006671858,0.0020461828,0.0019944948,0.008269877,0.12491545],"study_design_scores_gemma":[0.00009714309,0.0010900578,0.026734067,0.0019503228,0.0000727385,0.0009604582,0.9125612,0.0015898561,0.0015915703,0.007687943,0.045451336,0.00021325947],"about_ca_topic_score_codex":0.011595149,"about_ca_topic_score_gemma":0.030305455,"teacher_disagreement_score":0.9066355,"about_ca_system_score_codex":0.0048519555,"about_ca_system_score_gemma":0.008793323,"threshold_uncertainty_score":0.4937644},"labels":[],"label_agreement":null},{"id":"W2064653370","doi":"10.5555/2820518.2820527","title":"Co-evolution of infrastructure and source code: an empirical study","year":2015,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Server; Source code; Empirical research; Critical infrastructure; Database; Process (computing); Software engineering; Operating system; Programming language; Computer security","score_opus":0.02434610871695176,"score_gpt":0.31298362237117433,"score_spread":0.2886375136542226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064653370","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9982717,0.00017150806,0.0004851232,0.00011869745,0.000003946512,0.000047010693,0.00012835872,0.000012638669,0.0007609932],"genre_scores_gemma":[0.9988624,0.00008042626,0.0005007181,0.000032105017,0.0000069718217,0.000046794125,0.0002439431,0.000014957356,0.00021163041],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.977428,0.011702174,0.0018929503,0.0029060922,0.0047408114,0.0013299136],"domain_scores_gemma":[0.5008617,0.37992585,0.070991464,0.019274658,0.021189632,0.007756779],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019101763,0.000398162,0.00054407504,0.0052620205,0.0012973772,0.0027510345,0.0020821462,0.0016544028,0.0030230195],"category_scores_gemma":[0.15376776,0.00070226274,0.00065541675,0.0071286913,0.0032284479,0.005493152,0.0032011957,0.002946265,0.00079474435],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000060964587,0.00029696737,0.99325174,0.000030128296,0.000056557077,0.00014016232,0.002464255,0.00008750714,0.00007855341,0.00014069478,0.00012298881,0.0032694228],"study_design_scores_gemma":[0.000010960483,0.00027444493,0.9880415,0.00004826335,0.000060111128,0.00060220866,0.007756908,0.0019837827,0.00019214714,0.00016471493,0.0008479757,0.00001695205],"about_ca_topic_score_codex":0.006286184,"about_ca_topic_score_gemma":0.006984899,"teacher_disagreement_score":0.019101763,"about_ca_system_score_codex":0.0014785766,"about_ca_system_score_gemma":0.0016719532,"threshold_uncertainty_score":0.10102099},"labels":[],"label_agreement":null},{"id":"W2064790378","doi":"10.1109/icsm.2010.5609529","title":"Using clone detection to identify bugs in concurrent software","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Software bug; Thread (computing); clone (Java method); Software; Security bug; Programming language; Operating system; Biology; Software security assurance","score_opus":0.04002055785918808,"score_gpt":0.3544700996004586,"score_spread":0.31444954174127054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064790378","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19528721,0.00044743947,0.7994564,0.00017324189,0.00004582781,0.00016590272,0.000046317982,0.003423436,0.0009542254],"genre_scores_gemma":[0.7424765,0.00020930717,0.2558887,0.00011529953,0.00004870663,0.00008768423,0.000121101315,0.00025114155,0.0008015392],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9935529,0.0013758373,0.0005268731,0.0012896095,0.002938805,0.00031608602],"domain_scores_gemma":[0.93870115,0.039299898,0.00762266,0.0066544255,0.006777317,0.0009444997],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035098337,0.0010279675,0.0012287082,0.0048985137,0.00080248766,0.0023308096,0.0027978069,0.0018805229,0.00060522195],"category_scores_gemma":[0.03620809,0.00069866935,0.0010651224,0.001515888,0.0016214418,0.004454871,0.0017846038,0.001300335,0.00025346142],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053428387,0.00051648536,0.13623375,0.00043401917,0.0003352673,0.0025824886,0.0017296944,0.038853977,0.0907521,0.012512348,0.00086866564,0.7146469],"study_design_scores_gemma":[0.00015232952,0.0013083374,0.022463221,0.00015902288,0.0005143713,0.006679224,0.00035609741,0.789473,0.13659962,0.0378233,0.004160162,0.0003112457],"about_ca_topic_score_codex":0.002486318,"about_ca_topic_score_gemma":0.0020078716,"teacher_disagreement_score":0.0048985137,"about_ca_system_score_codex":0.0005302909,"about_ca_system_score_gemma":0.000909763,"threshold_uncertainty_score":0.018562019},"labels":[],"label_agreement":null},{"id":"W2065045338","doi":"10.1007/s11219-009-9081-z","title":"A new perspective on data homogeneity in software cost estimation: a study in the embedded systems domain","year":2009,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Türkiye Bilimsel ve Teknolojik Araştırma Kurumu","keywords":"Computer science; Software; Cost estimate; Domain (mathematical analysis); Estimation; Data mining; Empirical research; Perspective (graphical); Domain engineering; Software development; Cost driver; Software metric; Data science; Artificial intelligence; Software construction; Systems engineering; Engineering; Statistics","score_opus":0.08691323916499845,"score_gpt":0.4021292808525964,"score_spread":0.315216041687598,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2065045338","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07567355,0.0031333596,0.91064966,0.00439895,0.000112367445,0.00007041935,0.0002366654,0.00009900932,0.005625949],"genre_scores_gemma":[0.8362548,0.0022218276,0.15698422,0.0011605258,0.000996528,0.000121682104,0.00033583058,0.00026155036,0.0016630775],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.97665673,0.01343635,0.0014511057,0.0033594512,0.0044629625,0.00063339307],"domain_scores_gemma":[0.6015313,0.36060327,0.01100359,0.016022509,0.009178263,0.0016610553],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023308506,0.0009850585,0.002392342,0.0047776657,0.0014323463,0.0063779443,0.0041009984,0.0027530484,0.003452657],"category_scores_gemma":[0.1435379,0.001289496,0.002749367,0.0073946863,0.005193875,0.015940515,0.0047270497,0.0037895776,0.0002953201],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060091534,0.00044251667,0.049145866,0.0009823862,0.00076705636,0.0011944465,0.003748137,0.077098824,0.0040404405,0.7068805,0.0024355054,0.1526634],"study_design_scores_gemma":[0.00009664695,0.0006341134,0.017428517,0.00031281257,0.0005039385,0.0011951191,0.001780729,0.35717613,0.004465568,0.60754377,0.008707846,0.00015476916],"about_ca_topic_score_codex":0.0036582558,"about_ca_topic_score_gemma":0.0016158228,"teacher_disagreement_score":0.023308506,"about_ca_system_score_codex":0.0021811237,"about_ca_system_score_gemma":0.0014429237,"threshold_uncertainty_score":0.123268604},"labels":[],"label_agreement":null},{"id":"W2066455950","doi":"10.1109/tse.2015.2448531","title":"Assessing the Refactorability of Software Clones","year":2015,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; Computer science; clone (Java method); Maintainability; Programming language; Software maintenance; Code (set theory); Software evolution; Cloning (programming); Source code; Software; Software system; Software engineering; Set (abstract data type); Software construction; Biology; Genetics","score_opus":0.048312374942711414,"score_gpt":0.3005763212471097,"score_spread":0.2522639463043983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2066455950","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9134827,0.001171946,0.08142832,0.000092593546,0.000025979338,0.00017058078,0.0004584902,0.002222262,0.00094707834],"genre_scores_gemma":[0.9164401,0.00034313463,0.08077188,0.00004726043,0.000013490973,0.00008670095,0.00137397,0.00020849792,0.0007148846],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912062,0.0014448346,0.00087935536,0.0017790439,0.004350164,0.0003403075],"domain_scores_gemma":[0.92036873,0.041429427,0.01699581,0.0060111145,0.014050832,0.001144063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057318173,0.00095933134,0.0008038545,0.007010846,0.00052254327,0.0013133248,0.0013076955,0.0012564708,0.0006271232],"category_scores_gemma":[0.056809045,0.0004437759,0.00096720946,0.0022664517,0.0006983272,0.0019058272,0.0015550931,0.0007334761,0.00035344248],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004598708,0.00023409298,0.6067723,0.0010123986,0.00041835592,0.001008644,0.0020081713,0.024126701,0.04973201,0.0010218157,0.0010556033,0.31214988],"study_design_scores_gemma":[0.00008250737,0.0015417576,0.5830776,0.0004228523,0.0006935825,0.003089601,0.0017423896,0.31052265,0.088195756,0.003161649,0.007234302,0.00023538394],"about_ca_topic_score_codex":0.0034147662,"about_ca_topic_score_gemma":0.0042250105,"teacher_disagreement_score":0.007010846,"about_ca_system_score_codex":0.00070194196,"about_ca_system_score_gemma":0.0010225524,"threshold_uncertainty_score":0.030313134},"labels":[],"label_agreement":null},{"id":"W2066687109","doi":"10.1007/s11390-007-9051-5","title":"Developing Project Duration Models in Software Engineering","year":2007,"lang":"en","type":"article","venue":"Journal of Computer Science and Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Bell (Canada); École de Technologie Supérieure","funders":"","keywords":"Duration (music); Computer science; Benchmarking; Software; Software engineering; Theory of computation; Variable (mathematics); Software project management; Function point; Range (aeronautics); Industrial engineering; Operations research; Systems engineering; Software development; Software construction; Operating system; Engineering","score_opus":0.021558798790597647,"score_gpt":0.27681494181403127,"score_spread":0.25525614302343363,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2066687109","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017092602,0.00043818343,0.97832257,0.0003236033,0.000042210315,0.000065698434,0.00013098693,0.00021256087,0.0033716348],"genre_scores_gemma":[0.56683725,0.0017870793,0.4231686,0.00017413187,0.00013680846,0.0007002953,0.0008109233,0.00038352015,0.0060013966],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99728084,0.0014749545,0.00021288288,0.00029753998,0.00045833254,0.00027541368],"domain_scores_gemma":[0.9716912,0.023617577,0.0014834185,0.001021373,0.0016492467,0.0005372514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008868127,0.0012075158,0.0010730538,0.0021290348,0.00083567936,0.0025662929,0.0027724877,0.0018686826,0.0033957716],"category_scores_gemma":[0.03203012,0.0017231331,0.0016857891,0.0022743272,0.0009096838,0.0046831425,0.0016892981,0.002673353,0.0008801772],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000047073263,0.00009672209,0.0016069451,0.00010408104,0.00005366858,0.000067199195,0.00024696166,0.8449022,0.00036203966,0.11100046,0.0014418189,0.04007085],"study_design_scores_gemma":[0.000010033689,0.000020156038,0.00021917296,0.000030672767,0.000021972619,0.000012973617,0.000044457636,0.93657506,0.00015845336,0.0615319,0.0013649215,0.0000101580645],"about_ca_topic_score_codex":0.012476481,"about_ca_topic_score_gemma":0.010999805,"teacher_disagreement_score":0.012476481,"about_ca_system_score_codex":0.003188781,"about_ca_system_score_gemma":0.0033202036,"threshold_uncertainty_score":0.046899676},"labels":[],"label_agreement":null},{"id":"W2066687630","doi":"10.1007/s10664-013-9292-6","title":"Towards improving statistical modeling of software engineering data: think locally, act globally!","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Cluster analysis; Multivariate adaptive regression splines; Data mining; Machine learning; Software; Statistical model; Parametric statistics; Data modeling; Artificial intelligence; Data science; Regression analysis; Nonparametric regression; Mathematics; Statistics; Software engineering","score_opus":0.03219726578772242,"score_gpt":0.2902385702837676,"score_spread":0.2580413044960452,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2066687630","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042074416,0.0005568468,0.97375315,0.018446777,0.00016746964,0.00005312968,0.0001172218,0.0014743087,0.001223631],"genre_scores_gemma":[0.08202575,0.0010688298,0.90793806,0.005238143,0.0004739096,0.00025032272,0.00040459188,0.001254008,0.0013463143],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94334006,0.04207999,0.0018262382,0.0041256105,0.007812378,0.0008157023],"domain_scores_gemma":[0.7297781,0.1462482,0.012823354,0.081868425,0.025311468,0.003970454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07680857,0.00283248,0.0036348812,0.0043869806,0.0018545771,0.011362925,0.004437484,0.0054854727,0.0047256076],"category_scores_gemma":[0.24037509,0.002251046,0.0032083932,0.0046262937,0.007781678,0.034479827,0.008205913,0.015933225,0.0041406853],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029825518,0.0008558477,0.041392874,0.0012372265,0.0015482021,0.00018403976,0.00407631,0.07187728,0.008668546,0.24796177,0.043667864,0.5782317],"study_design_scores_gemma":[0.00007392813,0.00017875353,0.0036745905,0.00053571776,0.00021896724,0.00010608142,0.0016713551,0.24423127,0.0051847664,0.7248739,0.01909667,0.00015395989],"about_ca_topic_score_codex":0.007983583,"about_ca_topic_score_gemma":0.008212539,"teacher_disagreement_score":0.07680857,"about_ca_system_score_codex":0.002548395,"about_ca_system_score_gemma":0.008333705,"threshold_uncertainty_score":0.40620738},"labels":[],"label_agreement":null},{"id":"W2066998448","doi":"10.1145/1984701.1984706","title":"Measuring API documentation on the web","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":122,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Documentation; World Wide Web; Software documentation; Computer science; Social media; Software; Social software; Software development; Social web; Web 2.0; Web application; Web page; Software development process","score_opus":0.09346630203070735,"score_gpt":0.2630867740398661,"score_spread":0.16962047200915872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2066998448","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.989916,0.00051765767,0.001029224,0.000095684416,0.000011768716,0.00004014743,0.0010287077,0.00012598619,0.0072347997],"genre_scores_gemma":[0.99146384,0.0005162608,0.0031843723,0.000035424415,0.000046874982,0.000084855186,0.0027065328,0.00006569731,0.0018960106],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99345183,0.0017527698,0.0010179528,0.0005274247,0.0029436585,0.000306405],"domain_scores_gemma":[0.9045285,0.05454003,0.020007882,0.003754882,0.015300141,0.0018685447],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004019455,0.00032681754,0.00046785577,0.014934439,0.0006226173,0.0026366487,0.00043500998,0.0008461254,0.002026084],"category_scores_gemma":[0.051927734,0.000284474,0.00032354228,0.010005429,0.00042220028,0.003775347,0.001554205,0.0006181299,0.0011791015],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018153785,0.0002410775,0.88759226,0.00048946467,0.0001500011,0.0002187354,0.0045058993,0.0007724742,0.0027021384,0.00062066433,0.0026540067,0.09987176],"study_design_scores_gemma":[0.00000999222,0.00019332065,0.9785792,0.00017458074,0.0001066579,0.0006394426,0.004327496,0.0049037444,0.0024725113,0.00043794612,0.008111258,0.00004378113],"about_ca_topic_score_codex":0.0027390136,"about_ca_topic_score_gemma":0.003158168,"teacher_disagreement_score":0.014934439,"about_ca_system_score_codex":0.00048377065,"about_ca_system_score_gemma":0.00046628187,"threshold_uncertainty_score":0.021257162},"labels":[],"label_agreement":null},{"id":"W2067057008","doi":"10.1145/643603.643609","title":"Back to the future","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":119,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Aspect-oriented programming; Modular design; Modularity (biology); Computer science; Implementation; Redundancy (engineering); Extensibility; Separation of concerns; Programming language; Key (lock); Software engineering; Code (set theory); Software; Operating system","score_opus":0.012578341168893217,"score_gpt":0.2512104276890441,"score_spread":0.2386320865201509,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067057008","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004060543,0.026963297,0.0041009896,0.16286089,0.02286449,0.000033571807,0.0013581257,0.0005208893,0.77723724],"genre_scores_gemma":[0.06127376,0.028438766,0.004521835,0.057450756,0.005213146,0.00007136387,0.0015297855,0.0005804861,0.84092015],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999332,0.00013809257,0.00002713652,0.00013998845,0.00020776442,0.00015498248],"domain_scores_gemma":[0.9988249,0.00013124345,0.00007213927,0.00018964434,0.00044136835,0.0003406544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011602753,0.0005293288,0.0003114978,0.0006273049,0.0024485758,0.005832209,0.00064249535,0.002491707,0.18020973],"category_scores_gemma":[0.0043235836,0.00013618378,0.0003594144,0.00090201123,0.0018908044,0.008193677,0.002902893,0.0040433435,0.06805053],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000061284256,0.000027691822,0.0005823853,0.00014068969,0.000011548017,0.00015054354,0.00097351323,0.00012713697,0.00036500097,0.25765258,0.61820793,0.121699676],"study_design_scores_gemma":[0.0000012307206,0.0000041648445,0.00013246435,0.00006861218,0.0000011505499,0.000043750693,0.00016539723,0.000011203761,0.00003219879,0.0066230767,0.9929142,0.000002609026],"about_ca_topic_score_codex":0.0050203064,"about_ca_topic_score_gemma":0.0057280166,"teacher_disagreement_score":0.18020973,"about_ca_system_score_codex":0.002670188,"about_ca_system_score_gemma":0.0026988548,"threshold_uncertainty_score":0.6028616},"labels":[],"label_agreement":null},{"id":"W2067072439","doi":"10.1145/1370114.1370122","title":"Creating a cognitive metric of programming task difficulty","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Task (project management); Metric (unit); Selection (genetic algorithm); Software; Code (set theory); Human–computer interaction; Programming language; Cognition; Task analysis; Base (topology); Machine learning; Software engineering; Artificial intelligence; Psychology; Engineering","score_opus":0.02659999631309141,"score_gpt":0.27753081813807645,"score_spread":0.25093082182498505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067072439","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76089436,0.00037684006,0.21964593,0.0003290248,0.00026863674,0.0025688962,0.0020021063,0.0008572184,0.013056859],"genre_scores_gemma":[0.8910949,0.00011782232,0.099406205,0.00011142275,0.00007093229,0.0063087167,0.0017936638,0.00018449796,0.0009118895],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98795843,0.0052168826,0.0016053964,0.0016720806,0.0031180645,0.00042910632],"domain_scores_gemma":[0.8696207,0.09128313,0.01078131,0.009336494,0.015904434,0.0030739228],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011188008,0.0010279057,0.00081627164,0.004628211,0.0006173642,0.0026199939,0.0010478212,0.001051646,0.0024089874],"category_scores_gemma":[0.11462207,0.0004028474,0.00058113574,0.002138959,0.0011261097,0.0028691634,0.002440043,0.001504644,0.0004638335],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006577293,0.00616888,0.34989062,0.0028884076,0.001074406,0.00030111044,0.012441666,0.020718854,0.082625754,0.030892825,0.012230218,0.47418994],"study_design_scores_gemma":[0.0014606642,0.018647736,0.77047807,0.00039605712,0.00057138223,0.0005572362,0.004637939,0.0992204,0.04147776,0.036729336,0.025046233,0.00077708496],"about_ca_topic_score_codex":0.0010953721,"about_ca_topic_score_gemma":0.0012984644,"teacher_disagreement_score":0.011188008,"about_ca_system_score_codex":0.00093983195,"about_ca_system_score_gemma":0.0008289681,"threshold_uncertainty_score":0.059168577},"labels":[],"label_agreement":null},{"id":"W2067250880","doi":"10.5555/2820518.2820569","title":"Organizational volatility and post-release defects: a replication case study using data from Google Chrome","year":2015,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Outsourcing; Popularity; Volatility (finance); Business; Computer science; Directory; Replicate; World Wide Web; Knowledge management; Marketing; Finance; Operating system","score_opus":0.060518256411597054,"score_gpt":0.3147072335686175,"score_spread":0.2541889771570205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067250880","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9983199,0.000065348766,0.0006917576,0.000066235894,0.000004828394,0.000081809994,0.0004628222,0.000027292828,0.00027995848],"genre_scores_gemma":[0.9963921,0.000047968522,0.0019701854,0.00003771475,0.000013839199,0.00012336436,0.0011033025,0.000021939526,0.00028945308],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9909151,0.0047658063,0.00082910585,0.0014349901,0.0015451835,0.00050975324],"domain_scores_gemma":[0.90250343,0.054310214,0.015059375,0.015594817,0.010884887,0.0016472713],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010282141,0.0005558202,0.0006591169,0.00355258,0.0012536842,0.0015292554,0.002075019,0.0017000554,0.0008709373],"category_scores_gemma":[0.051926654,0.0004851667,0.0011922937,0.0043948977,0.0013041538,0.001680185,0.0013786611,0.0013928548,0.0004046446],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058211817,0.0015956595,0.9559567,0.0002614817,0.0003656178,0.003179928,0.013209667,0.002551248,0.0020728717,0.00052739395,0.0015693677,0.018127955],"study_design_scores_gemma":[0.00019154155,0.0014426775,0.95891935,0.000100358826,0.00027408605,0.0018875883,0.017628497,0.011760616,0.003226514,0.0006476085,0.0037767484,0.00014441309],"about_ca_topic_score_codex":0.026591904,"about_ca_topic_score_gemma":0.026666792,"teacher_disagreement_score":0.98971784,"about_ca_system_score_codex":0.0014704914,"about_ca_system_score_gemma":0.0011633792,"threshold_uncertainty_score":0.054377854},"labels":[],"label_agreement":null},{"id":"W2067377566","doi":"10.1016/j.scico.2010.11.010","title":"An empirical study on inconsistent changes to code clones at the release level","year":2010,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software evolution; Software release life cycle; Software; clone (Java method); Software quality; Software development; Perspective (graphical); Quality (philosophy); Code (set theory); Empirical research; Source code; Software engineering; Software construction; Programming language; Artificial intelligence; Statistics; Biology","score_opus":0.06513932518415562,"score_gpt":0.3623689497781702,"score_spread":0.2972296245940146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067377566","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9979311,0.00015535466,0.00059378176,0.000118797805,0.000006997066,0.00003189968,0.00009503852,0.0000145375225,0.0010524988],"genre_scores_gemma":[0.99855787,0.00007305897,0.00060099136,0.00008866614,0.000010867679,0.000027032824,0.00022662127,0.000027627002,0.0003872698],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98147994,0.009311834,0.0018945013,0.00202269,0.0046020513,0.0006890274],"domain_scores_gemma":[0.2937554,0.5736138,0.08197597,0.023026863,0.02422729,0.0034006434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01541674,0.0003734003,0.00044249868,0.0025848215,0.0012706098,0.0019593318,0.0016930809,0.0017378067,0.0039970726],"category_scores_gemma":[0.28005132,0.0005431284,0.00041070816,0.002778576,0.0021202494,0.0034338862,0.0013837742,0.0029426112,0.0006066738],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012668489,0.0023196468,0.9624238,0.0002080122,0.00014403474,0.0004646898,0.00793324,0.0005256055,0.0013498731,0.0010989787,0.00063783623,0.021627558],"study_design_scores_gemma":[0.00012782718,0.0017676785,0.97916585,0.00013384921,0.00021180559,0.0011174507,0.009718454,0.0029676908,0.0016309347,0.0009333957,0.002181905,0.00004305382],"about_ca_topic_score_codex":0.0037098727,"about_ca_topic_score_gemma":0.0048518726,"teacher_disagreement_score":0.01541674,"about_ca_system_score_codex":0.0012194152,"about_ca_system_score_gemma":0.001300657,"threshold_uncertainty_score":0.08153248},"labels":[],"label_agreement":null},{"id":"W2067490448","doi":"","title":"The Impact of Mislabelling on the Performance and Interpretation of Defect Prediction Models","year":2018,"lang":"en","type":"article","venue":"Institutional Repositories DataBase (IRDB)","topic":"Software Engineering Research","field":"Computer Science","cited_by":89,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Noise (video); Reliability (semiconductor); Interpretation (philosophy); Artificial intelligence; Predictive modelling; Machine learning; Rank (graph theory); Training set; Data modeling; Recall; Data mining; Database; Mathematics; Cognitive psychology","score_opus":0.01975086518738715,"score_gpt":0.2721726546332785,"score_spread":0.25242178944589133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067490448","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8833206,0.0139223235,0.076873966,0.006097048,0.002082954,0.00023742742,0.004001985,0.007510995,0.0059526614],"genre_scores_gemma":[0.9622378,0.00069331966,0.02882803,0.0011770495,0.0002764061,0.00004801349,0.0044941814,0.0008810057,0.0013641263],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.93393534,0.039996423,0.004854595,0.012092268,0.007766059,0.0013552627],"domain_scores_gemma":[0.4727199,0.47783893,0.010807097,0.020585205,0.015147119,0.0029016547],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06652968,0.0029524153,0.0028864462,0.0036974417,0.0025760452,0.008095733,0.0033605082,0.0052623525,0.0014392135],"category_scores_gemma":[0.27014148,0.0009846757,0.0020166335,0.0025833002,0.0022184956,0.0062791845,0.003598867,0.005343459,0.0014481546],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010473274,0.001947974,0.286235,0.0015326933,0.0032655573,0.0011825587,0.0023265707,0.16331165,0.012244562,0.0033391358,0.035282746,0.4788583],"study_design_scores_gemma":[0.0003763012,0.0012328696,0.039852202,0.0005927336,0.0014568744,0.0013064332,0.0015080447,0.91552204,0.017078247,0.015624422,0.005165423,0.00028439835],"about_ca_topic_score_codex":0.015242981,"about_ca_topic_score_gemma":0.01657603,"teacher_disagreement_score":0.06652968,"about_ca_system_score_codex":0.0029061988,"about_ca_system_score_gemma":0.0030036774,"threshold_uncertainty_score":0.35184675},"labels":[],"label_agreement":null},{"id":"W2067948082","doi":"10.1109/ictai.2006.70","title":"Intelligent Software Measurement System for Automating the Goal-Question-Metrics Process","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Software measurement; Software engineering; Process (computing); Knowledge base; Software; Ontology; Software system; Software construction; Artificial intelligence; Programming language","score_opus":0.0279658557052207,"score_gpt":0.28054867202434036,"score_spread":0.25258281631911966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067948082","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012026712,0.00010468119,0.9534196,0.0002405153,0.000058027905,0.00047942429,0.00021585698,0.028465196,0.00499],"genre_scores_gemma":[0.12118054,0.00008916973,0.87389445,0.00014822237,0.000029825424,0.0006080738,0.0007970998,0.00032979515,0.0029229238],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99778533,0.0006558814,0.00027020994,0.0004508026,0.0007446135,0.00009317273],"domain_scores_gemma":[0.994118,0.0020198673,0.00072256016,0.0008018441,0.002042576,0.00029516147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043635005,0.0006446498,0.000869483,0.001792702,0.00077557575,0.0014257921,0.0013359505,0.00083438674,0.0035093066],"category_scores_gemma":[0.00989912,0.00042729362,0.00044604423,0.0011442886,0.00058830425,0.002375523,0.0012037571,0.0011795964,0.0016988752],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005853538,0.000782274,0.008304673,0.000619652,0.00018329732,0.0002873421,0.0015303815,0.019474821,0.04593307,0.058305625,0.023447428,0.8405461],"study_design_scores_gemma":[0.00033799963,0.00093264104,0.009146391,0.00020411736,0.0003440546,0.0005194608,0.00033248094,0.7773637,0.07300128,0.023894574,0.11372157,0.00020173758],"about_ca_topic_score_codex":0.0025991895,"about_ca_topic_score_gemma":0.0024735357,"teacher_disagreement_score":0.0043635005,"about_ca_system_score_codex":0.0013534797,"about_ca_system_score_gemma":0.0030861925,"threshold_uncertainty_score":0.023076713},"labels":[],"label_agreement":null},{"id":"W2068001199","doi":"10.1109/icsm.2013.24","title":"Combining Static and Dynamic Analyses to Reverse-Engineer Scenario Diagrams","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Sequence diagram; Computer science; Reverse engineering; Static analysis; Unified Modeling Language; Source code; Java; Programming language; Class diagram; Program comprehension; Dynamic program analysis; Instrumentation (computer programming); Use Case Diagram; Static program analysis; Activity diagram; UML tool; Applications of UML; Software engineering; Software; Software system; Software development","score_opus":0.02648140674206771,"score_gpt":0.3058633727790844,"score_spread":0.27938196603701665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068001199","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004199403,0.00004005967,0.99257267,0.00007540047,0.00001897779,0.000074745876,0.000061504514,0.0019182019,0.0010390134],"genre_scores_gemma":[0.076974176,0.00018041542,0.91934776,0.00007755627,0.000021349586,0.0001301108,0.00048276925,0.0013086655,0.0014772083],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946785,0.0018318328,0.00036638803,0.0006188465,0.0022114262,0.000293007],"domain_scores_gemma":[0.983841,0.008609004,0.001099892,0.0035561759,0.0027162551,0.00017768463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004571468,0.001816784,0.00066014635,0.00474131,0.0007354064,0.002516142,0.0009775823,0.0010436156,0.0037985244],"category_scores_gemma":[0.018693078,0.00092497235,0.0020608234,0.0012241366,0.001024444,0.0031599023,0.0022981009,0.0019467652,0.0018995517],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019704846,0.00034156648,0.0075531974,0.0008584938,0.00028184924,0.0017380139,0.0024558876,0.12724584,0.105398886,0.11250722,0.0042089317,0.637213],"study_design_scores_gemma":[0.000058089645,0.00023284972,0.002291761,0.00043951336,0.0002802825,0.001576884,0.0006613216,0.6206654,0.15746118,0.12780263,0.08831919,0.00021083497],"about_ca_topic_score_codex":0.0024857041,"about_ca_topic_score_gemma":0.0043550464,"teacher_disagreement_score":0.00474131,"about_ca_system_score_codex":0.0008821603,"about_ca_system_score_gemma":0.002172365,"threshold_uncertainty_score":0.024176538},"labels":[],"label_agreement":null},{"id":"W2068059071","doi":"10.1145/1985404.1985423","title":"Towards flexible code clone detection, management, and refactoring in IDE","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; Computer science; clone (Java method); Programming language; Software maintenance; Code (set theory); Plug-in; Eclipse; Source code; Java; Software engineering; Software development; Software; Set (abstract data type)","score_opus":0.05012460558356874,"score_gpt":0.27309649654655993,"score_spread":0.2229718909629912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068059071","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033554606,0.00016712611,0.94855535,0.00011556599,0.000024815388,0.00009178888,0.00006031713,0.016941704,0.00048879784],"genre_scores_gemma":[0.10987959,0.000121941936,0.8871137,0.00010837573,0.000021145677,0.00009460969,0.00037399202,0.0009324329,0.0013543039],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99675757,0.00078573945,0.00035471455,0.0007119039,0.0011828025,0.00020737153],"domain_scores_gemma":[0.98853636,0.0040624924,0.0011606304,0.0036572772,0.0022429544,0.00034036132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004637732,0.00065071206,0.0009210574,0.0014075543,0.0005505516,0.0016486942,0.0022978003,0.0012795391,0.00065402745],"category_scores_gemma":[0.012933055,0.00069671153,0.00062540243,0.00072490104,0.0006416212,0.0027967866,0.0018383267,0.0019030265,0.00081883545],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045721323,0.00032999268,0.013249601,0.00023433058,0.00008369582,0.00047077864,0.00087625644,0.013916782,0.12594752,0.0056566168,0.0035865472,0.83519065],"study_design_scores_gemma":[0.0002752802,0.0005928845,0.0083424095,0.00013070977,0.00013756866,0.0017331797,0.00028949743,0.6110433,0.3314289,0.016508117,0.02933803,0.00018019945],"about_ca_topic_score_codex":0.001107108,"about_ca_topic_score_gemma":0.0013433818,"teacher_disagreement_score":0.004637732,"about_ca_system_score_codex":0.00040306768,"about_ca_system_score_gemma":0.0011197068,"threshold_uncertainty_score":0.024526954},"labels":[],"label_agreement":null},{"id":"W2068204598","doi":"10.5555/2667089.2667097","title":"Revisiting bug triage and resolution practices","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Triage; Computer science; Process (computing); Categorization; Software bug; Computer security; Data science; Software; Artificial intelligence; Medical emergency; Medicine","score_opus":0.055471890546457715,"score_gpt":0.33548610964797404,"score_spread":0.2800142191015163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068204598","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.952837,0.003757314,0.022800414,0.009225894,0.00025829306,0.00081998005,0.00024257817,0.0011425589,0.008916065],"genre_scores_gemma":[0.97708225,0.0013078052,0.018688062,0.00042778428,0.00005733357,0.00038377102,0.00025784454,0.00018732327,0.0016077213],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9451767,0.028961057,0.006056403,0.0055247797,0.012269452,0.0020114377],"domain_scores_gemma":[0.5902565,0.24496244,0.06660965,0.027092956,0.065735675,0.005342789],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.062475022,0.0006214648,0.00061191333,0.011301861,0.004209407,0.007462978,0.0028326078,0.001827521,0.0024076172],"category_scores_gemma":[0.29920006,0.0012646816,0.0006627887,0.0065445676,0.0034580773,0.008661898,0.004622051,0.0032820767,0.0006349128],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022226419,0.00062876893,0.19841552,0.0016293122,0.00011818421,0.00090026367,0.34781894,0.0010051609,0.008029426,0.0027994886,0.008860905,0.42957172],"study_design_scores_gemma":[0.00015602319,0.0017466455,0.4571789,0.0049189716,0.00025443805,0.0027502251,0.41516048,0.012383009,0.010856076,0.0046961643,0.0894657,0.00043338205],"about_ca_topic_score_codex":0.016244646,"about_ca_topic_score_gemma":0.016115284,"teacher_disagreement_score":0.062475022,"about_ca_system_score_codex":0.009133414,"about_ca_system_score_gemma":0.010702241,"threshold_uncertainty_score":0.33040345},"labels":[],"label_agreement":null},{"id":"W2068391864","doi":"10.3138/jsp.43.2.188","title":"Automated Document Analyser for Screening of Journal Articles","year":2011,"lang":"en","type":"article","venue":"Journal of Scholarly Publishing","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Analyser; Acknowledgement; Computer science; Macro; Visual Basic for Applications; Task (project management); Information retrieval; Word processing; Software; Data science; World Wide Web; Natural language processing; Programming language; Computer security; Management","score_opus":0.07033598059158333,"score_gpt":0.2975578062188363,"score_spread":0.227221825627253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068391864","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019953838,0.0015329415,0.661288,0.00071443414,0.00093696965,0.0058696475,0.02220069,0.27373013,0.013773355],"genre_scores_gemma":[0.029423945,0.00066657073,0.9296655,0.00029061423,0.00046172226,0.004927152,0.013489843,0.010231104,0.010843534],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9823267,0.004961415,0.0043470375,0.002930689,0.004882223,0.00055194384],"domain_scores_gemma":[0.8811267,0.06150664,0.00866117,0.0096152695,0.037513953,0.0015762679],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.018341972,0.002346491,0.002317631,0.014861526,0.0016781306,0.0046346053,0.0017496346,0.0013309581,0.033343315],"category_scores_gemma":[0.068904914,0.0011540387,0.0013427869,0.0077095036,0.0009203794,0.0036249391,0.0019723244,0.0017263573,0.03539443],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020548631,0.0002881874,0.0048444783,0.0034012294,0.00024951462,0.0009764571,0.0018400626,0.0012671507,0.08344672,0.0048696166,0.13385794,0.7629038],"study_design_scores_gemma":[0.00090798363,0.0013426002,0.037633687,0.0013323196,0.00050879084,0.004460278,0.0017177123,0.06838022,0.3049565,0.008410901,0.5693791,0.00096991105],"about_ca_topic_score_codex":0.0016089535,"about_ca_topic_score_gemma":0.0013767142,"teacher_disagreement_score":0.9953654,"about_ca_system_score_codex":0.0014537541,"about_ca_system_score_gemma":0.0039744815,"threshold_uncertainty_score":0.11154449},"labels":[],"label_agreement":null},{"id":"W2068919549","doi":"10.1016/j.jocs.2010.12.002","title":"Examining random and designed tests to detect code mistakes in scientific software","year":2010,"lang":"en","type":"article","venue":"Journal of Computational Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Royal Military College of Canada","funders":"","keywords":"Computer science; Mistake; Code (set theory); Test (biology); Software; Code coverage; Software bug; Programming language; Reliability engineering; Engineering","score_opus":0.029649452150085878,"score_gpt":0.29610406219953816,"score_spread":0.2664546100494523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068919549","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96321505,0.00015756782,0.033827852,0.00010653798,0.00006992443,0.00013205547,0.00014852722,0.0012335089,0.0011089413],"genre_scores_gemma":[0.97174925,0.000028382368,0.02707299,0.00012870734,0.000012913884,0.0000830152,0.00017742398,0.00015653149,0.0005907516],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98499227,0.0072634793,0.0014005047,0.0023655798,0.0034965875,0.0004816998],"domain_scores_gemma":[0.56678104,0.35375562,0.031259812,0.017517947,0.028271228,0.0024142396],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008245439,0.00079264864,0.00046620058,0.0020642607,0.00053193135,0.0010914291,0.0018079447,0.0017905248,0.0010580014],"category_scores_gemma":[0.20806697,0.000550252,0.0005419822,0.00096033,0.0010032696,0.0015467412,0.00089221715,0.0010512348,0.0003681114],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005283432,0.0026990357,0.6421588,0.0011020895,0.0007906306,0.0020973755,0.005067284,0.027779551,0.04915625,0.0038903083,0.0049625942,0.25501263],"study_design_scores_gemma":[0.0009639808,0.01437722,0.37110934,0.00051375054,0.0010072447,0.0038933733,0.0034024834,0.43728274,0.1493418,0.011352084,0.0064092726,0.00034680052],"about_ca_topic_score_codex":0.0017631366,"about_ca_topic_score_gemma":0.0034729245,"teacher_disagreement_score":0.99175453,"about_ca_system_score_codex":0.0009844356,"about_ca_system_score_gemma":0.0017094746,"threshold_uncertainty_score":0.04360658},"labels":[],"label_agreement":null},{"id":"W2069822733","doi":"10.1145/1088622.1088664","title":"Identification of question types and answer types for an explanation component in software release planning","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Component (thermodynamics); Identification (biology); Computer science; Software; Software system; Software engineering; Component-based software engineering; Data science; Programming language","score_opus":0.025248649677470567,"score_gpt":0.3087529233994565,"score_spread":0.2835042737219859,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2069822733","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023507876,0.00025192474,0.96820104,0.00079641235,0.000047970392,0.0008951299,0.00019030082,0.003014568,0.003094711],"genre_scores_gemma":[0.16346839,0.00025975067,0.8294123,0.0003081684,0.00006232196,0.00095558434,0.000798782,0.00067381683,0.004060815],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9707252,0.019166844,0.002822391,0.002392951,0.003807074,0.0010855354],"domain_scores_gemma":[0.9027149,0.07825581,0.0047800634,0.005490212,0.0076309354,0.001128062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026857963,0.0014703885,0.001392278,0.0034118725,0.0016634864,0.0043729157,0.0018596931,0.0037935309,0.005207105],"category_scores_gemma":[0.06069608,0.0016114594,0.0017557014,0.00154715,0.0029988452,0.009301682,0.003883161,0.0033188644,0.0020527323],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036803547,0.0007585362,0.017917404,0.002518933,0.0001470442,0.0016369725,0.03324365,0.008825606,0.04183463,0.11544452,0.006560332,0.76743203],"study_design_scores_gemma":[0.0008974654,0.0023464658,0.028261019,0.0029053693,0.0011529669,0.005328289,0.014824398,0.2700953,0.24163818,0.19931856,0.23220523,0.0010266991],"about_ca_topic_score_codex":0.0019483615,"about_ca_topic_score_gemma":0.0016936412,"teacher_disagreement_score":0.026857963,"about_ca_system_score_codex":0.0014853058,"about_ca_system_score_gemma":0.0025781868,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2069963480","doi":"10.1007/s11219-005-4253-y","title":"Virtual Software Engineering Laboratories in Support of Trade-off Analyses","year":2005,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Bundesministerium für Bildung und Forschung","keywords":"Product (mathematics); Computer science; Software engineering; Software; Quality (philosophy); Software quality; Systems engineering; Empirical research; Software quality control; Engineering; Software development; Process management; Risk analysis (engineering); Business; Operating system","score_opus":0.04235531741497253,"score_gpt":0.3502672804459467,"score_spread":0.3079119630309741,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2069963480","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35214695,0.0003564356,0.5999127,0.0012455212,0.00021873666,0.00023264285,0.00023429627,0.0110116955,0.03464103],"genre_scores_gemma":[0.92695373,0.00004803187,0.070484824,0.000053153763,0.000053617267,0.000103432016,0.00013483159,0.00018652437,0.0019818288],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99434245,0.0035422125,0.00018698975,0.000559806,0.0010066483,0.00036184493],"domain_scores_gemma":[0.943393,0.03302082,0.003345401,0.011878735,0.004930351,0.0034317658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008525589,0.0006364626,0.00075634633,0.002085861,0.0012129009,0.0045494833,0.003464938,0.0014429181,0.010745098],"category_scores_gemma":[0.03770426,0.0005367282,0.0003035256,0.0017057372,0.0009963528,0.005674752,0.0048225857,0.0012650603,0.0014301178],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010404404,0.0024856133,0.029309237,0.00025126943,0.0002182032,0.0009707464,0.0017154763,0.18887661,0.022949431,0.12215649,0.013294006,0.6073685],"study_design_scores_gemma":[0.0010205741,0.0010971609,0.004030733,0.00005437948,0.00013828246,0.0003703275,0.00053077197,0.9105879,0.016724626,0.051912446,0.013443767,0.00008902187],"about_ca_topic_score_codex":0.0008143728,"about_ca_topic_score_gemma":0.0010707853,"teacher_disagreement_score":0.010745098,"about_ca_system_score_codex":0.0011046577,"about_ca_system_score_gemma":0.002080726,"threshold_uncertainty_score":0.045088172},"labels":[],"label_agreement":null},{"id":"W2070321219","doi":"10.1145/1287624.1287673","title":"Does a programmer's activity indicate knowledge of code?","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":95,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Center for Advanced Study, University of Illinois at Urbana-Champaign; Natural Sciences and Engineering Research Council of Canada","keywords":"Programmer; Computer science; Programming language; Code (set theory); Java; Base (topology); Task (project management); Knowledge base; Software engineering; World Wide Web","score_opus":0.02103923281343554,"score_gpt":0.3109258293373323,"score_spread":0.28988659652389676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2070321219","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9911238,0.00019303057,0.004244436,0.0005972275,0.000008321082,0.000017172251,0.00014626877,0.000037719252,0.0036320637],"genre_scores_gemma":[0.99825245,0.00010485326,0.0012571511,0.00004236879,0.0000084991625,0.000010292641,0.000096295036,0.0000071357344,0.00022102546],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99567246,0.0023819313,0.00034332337,0.0005717318,0.0007959391,0.00023454835],"domain_scores_gemma":[0.8112066,0.13214068,0.036961585,0.007857656,0.008053681,0.003779785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004193515,0.0002371794,0.00032779548,0.002495353,0.00031493747,0.0015039258,0.00059192127,0.0014139828,0.0016367566],"category_scores_gemma":[0.07863639,0.0003680154,0.00031566704,0.0018018589,0.0010561218,0.0034608028,0.0008913875,0.0007183001,0.00047628654],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000106573054,0.00008706871,0.98010397,0.00007467665,0.00004372223,0.000055713554,0.003452974,0.00038522118,0.00067426875,0.0002671447,0.00015380277,0.014594962],"study_design_scores_gemma":[0.0000120401,0.00016651195,0.9912503,0.000042120933,0.000025256695,0.0002464492,0.002736685,0.003729545,0.00037528735,0.0007506495,0.0006466666,0.000018366683],"about_ca_topic_score_codex":0.003103468,"about_ca_topic_score_gemma":0.0049441205,"teacher_disagreement_score":0.004193515,"about_ca_system_score_codex":0.0005763851,"about_ca_system_score_gemma":0.0005258567,"threshold_uncertainty_score":0.022177696},"labels":[],"label_agreement":null},{"id":"W2070515606","doi":"10.1109/swan.2015.7070488","title":"MARFCAT: Fast code analysis for defects and vulnerabilities","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Bytecode; NIST; Source code; Static program analysis; Static analysis; Programming language; Open source; Code (set theory); Reverse engineering; Software; Natural language processing; Software development; Java","score_opus":0.04100671325540597,"score_gpt":0.2914485342778601,"score_spread":0.25044182102245416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2070515606","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03786512,0.00076153595,0.7271151,0.00029136025,0.00016998273,0.00032607425,0.006674301,0.22385359,0.0029430485],"genre_scores_gemma":[0.1870531,0.0002514496,0.78955704,0.00021511644,0.00012219592,0.00033596208,0.012853044,0.0045846077,0.0050275656],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99762803,0.00024323328,0.00010378062,0.0005623166,0.0012672037,0.00019545517],"domain_scores_gemma":[0.9938712,0.002291384,0.0008976792,0.0013766021,0.0013681524,0.00019483599],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014743642,0.002065028,0.0010075575,0.006894246,0.00072505034,0.0013497504,0.0026820977,0.001694508,0.0056741303],"category_scores_gemma":[0.009214379,0.0007726199,0.0014084384,0.002338134,0.0006433658,0.0024040067,0.0014322686,0.0014858461,0.004969962],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047519014,0.00034675293,0.02247997,0.0005243621,0.00032748136,0.0004960369,0.00024088286,0.03336967,0.03364149,0.0057229055,0.09119324,0.81118196],"study_design_scores_gemma":[0.00010063198,0.00039767995,0.012116028,0.00009807124,0.00007781741,0.0012376901,0.00007370978,0.871341,0.06497786,0.009746756,0.039660536,0.00017234574],"about_ca_topic_score_codex":0.005521483,"about_ca_topic_score_gemma":0.008301211,"teacher_disagreement_score":0.006894246,"about_ca_system_score_codex":0.00077324634,"about_ca_system_score_gemma":0.001510584,"threshold_uncertainty_score":0.018981814},"labels":[],"label_agreement":null},{"id":"W2070548384","doi":"10.1109/msr.2013.6624005","title":"Deficient documentation detection a methodology to locate deficient project documentation using topic analysis","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Documentation; Python (programming language); Computer science; World Wide Web; Scope (computer science); Internal documentation; Application programming interface; Software engineering; Programming language; Information retrieval; Software; Software development","score_opus":0.09632895295661605,"score_gpt":0.37378885297227377,"score_spread":0.27745990001565773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2070548384","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12883142,0.0012424688,0.8502468,0.00057259784,0.00025310554,0.002133173,0.005141729,0.007250115,0.00432855],"genre_scores_gemma":[0.23183848,0.0003906896,0.7540273,0.00014111931,0.00014378384,0.002992497,0.0065534883,0.0005014346,0.0034111897],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9893457,0.0028381716,0.0017776261,0.0030005153,0.0024540073,0.0005840494],"domain_scores_gemma":[0.9554649,0.021614272,0.007690477,0.0039060144,0.010340579,0.0009838771],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012293177,0.0013175883,0.0012885149,0.019518098,0.0020233411,0.0031260417,0.0019477644,0.0013935694,0.0023077982],"category_scores_gemma":[0.03765259,0.0009598964,0.0013852991,0.011312373,0.00088570407,0.003475566,0.003471786,0.001836622,0.0020002439],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007299774,0.00066083914,0.09749536,0.0022058382,0.00046201947,0.00078335986,0.022170745,0.003368313,0.04266163,0.008873029,0.027800286,0.7927886],"study_design_scores_gemma":[0.0005790488,0.0010889015,0.21556583,0.0012746901,0.0013622595,0.004438809,0.034889307,0.3869078,0.11994859,0.049106065,0.18387403,0.000964729],"about_ca_topic_score_codex":0.006043989,"about_ca_topic_score_gemma":0.0072632832,"teacher_disagreement_score":0.98770684,"about_ca_system_score_codex":0.0010615451,"about_ca_system_score_gemma":0.0039214953,"threshold_uncertainty_score":0.06501329},"labels":[],"label_agreement":null},{"id":"W2070873282","doi":"10.1007/s10664-011-9169-5","title":"The evolution of Java build systems","year":2011,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Java; Executable; Software evolution; Software maintenance; Software system; Software engineering; Source code; Software development; Legacy system; Overhead (engineering); Codebase; Source lines of code; Software; Operating system; Distributed computing; Software construction","score_opus":0.028469191643593714,"score_gpt":0.2551634501681812,"score_spread":0.2266942585245875,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2070873282","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9871273,0.00026796685,0.0031248978,0.00075774675,0.000011773496,0.000015331463,0.00019734964,0.00006490428,0.008432845],"genre_scores_gemma":[0.99679923,0.00011045835,0.0015857483,0.000040421888,0.0000054261604,0.0000042842835,0.00017535474,0.000025642823,0.0012532683],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99857044,0.0005606273,0.00007959474,0.00018119194,0.00046843308,0.00013970204],"domain_scores_gemma":[0.9809176,0.010537713,0.0025044386,0.0019884747,0.0034304562,0.00062136014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026556815,0.00015917062,0.00011203029,0.0014071802,0.00048589584,0.0018382899,0.00050928973,0.0006459558,0.0020588543],"category_scores_gemma":[0.036080834,0.0002888966,0.00019704274,0.0013632404,0.0008112526,0.0027978243,0.0007468712,0.0009841537,0.0002590135],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004377382,0.00053932786,0.66513884,0.0001742917,0.00013936554,0.00051617314,0.0048564975,0.03371859,0.008830003,0.046881385,0.0030188067,0.23574893],"study_design_scores_gemma":[0.000034920013,0.00022864016,0.8549366,0.00009024659,0.000095599156,0.00050848554,0.0027887698,0.10220103,0.0047083967,0.017986499,0.01637227,0.000048649574],"about_ca_topic_score_codex":0.012510412,"about_ca_topic_score_gemma":0.015489274,"teacher_disagreement_score":0.012510412,"about_ca_system_score_codex":0.0019933917,"about_ca_system_score_gemma":0.0009478204,"threshold_uncertainty_score":0.024875224},"labels":[],"label_agreement":null},{"id":"W2071194128","doi":"10.1109/saner.2015.7081842","title":"The influence of App churn on App success and StackOverflow discussions","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Android (operating system); Mobile apps; Computer science; Documentation; World Wide Web; Android app; Smartphone app; App store; Software; Operating system","score_opus":0.017387847514773912,"score_gpt":0.2761783646854437,"score_spread":0.2587905171706698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2071194128","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9981407,0.00011465519,0.00032001353,0.00006746765,0.0000058836254,0.00001487325,0.00009156369,0.000029498982,0.0012152966],"genre_scores_gemma":[0.99868685,0.000067518275,0.0003561129,0.000019792458,0.000021421698,0.00002131957,0.00019167515,0.000021546908,0.0006138281],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9932092,0.0022686773,0.0006784484,0.0008620126,0.0023180146,0.00066361],"domain_scores_gemma":[0.7703589,0.15380265,0.051681455,0.005904637,0.0121743325,0.006077979],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064747743,0.00047282808,0.00048469426,0.0028318157,0.00096569065,0.0021103956,0.0004948777,0.00056257506,0.0018976344],"category_scores_gemma":[0.06884083,0.0003681905,0.0004917678,0.0014974745,0.0010191153,0.0022836514,0.0016429495,0.0013861503,0.0006413443],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019573861,0.0001005846,0.9824867,0.00004855736,0.000070407994,0.00015367054,0.0033564845,0.00019316304,0.0007103303,0.00008611961,0.00030796826,0.012290339],"study_design_scores_gemma":[0.0000023048706,0.00009188388,0.9971021,0.000014470873,0.000021615066,0.000108793414,0.0009923024,0.0009857424,0.00029305337,0.00005643224,0.00031843362,0.000012846306],"about_ca_topic_score_codex":0.0049484344,"about_ca_topic_score_gemma":0.007696299,"teacher_disagreement_score":0.0064747743,"about_ca_system_score_codex":0.0006372611,"about_ca_system_score_gemma":0.00062366104,"threshold_uncertainty_score":0.034242332},"labels":[],"label_agreement":null},{"id":"W2071255138","doi":"10.1007/s10664-014-9323-y","title":"Recommending reference API documentation","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Documentation; Computer science; Categorization; Programmer; Set (abstract data type); Internal documentation; Information retrieval; World Wide Web; Software; Artificial intelligence; Programming language; Software system","score_opus":0.03239301421113391,"score_gpt":0.30504308565920896,"score_spread":0.27265007144807507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2071255138","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27092254,0.0065027867,0.5602111,0.0067295046,0.0027751254,0.0014860531,0.007873271,0.0626672,0.08083245],"genre_scores_gemma":[0.5664424,0.0028332015,0.36205223,0.0011645468,0.0006287945,0.00052105705,0.017681992,0.0036037862,0.04507205],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946477,0.0015764543,0.00044806566,0.0006753393,0.0024434072,0.00020898004],"domain_scores_gemma":[0.96956563,0.007850987,0.0014458302,0.0057420204,0.014515078,0.0008804859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003378434,0.0011279515,0.0010793498,0.010040276,0.001604274,0.002832006,0.0015979406,0.002183166,0.014816348],"category_scores_gemma":[0.06631322,0.000676202,0.0010086986,0.005870799,0.00034013236,0.004117116,0.0014610904,0.0017994692,0.00935573],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034934568,0.00060676876,0.016954489,0.00067614333,0.00010675158,0.00040169773,0.00048376885,0.005559518,0.00857547,0.004246269,0.14295493,0.8190849],"study_design_scores_gemma":[0.0004255502,0.001383113,0.037759885,0.0013473724,0.0008493579,0.002178494,0.002091787,0.5839574,0.048242595,0.027817177,0.29359892,0.00034850414],"about_ca_topic_score_codex":0.009660594,"about_ca_topic_score_gemma":0.01985606,"teacher_disagreement_score":0.014816348,"about_ca_system_score_codex":0.00097325904,"about_ca_system_score_gemma":0.0031696665,"threshold_uncertainty_score":0.049565613},"labels":[],"label_agreement":null},{"id":"W2071311708","doi":"10.1145/1297846.1297904","title":"Comprehending implementation recipes of framework-provided concepts through dynamic analysis","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Slicing; Program comprehension; Software engineering; Cluster analysis; Program slicing; Data science; Usability; Data mining; Programming language; World Wide Web; Human–computer interaction; Artificial intelligence; Software; Software system","score_opus":0.02548012517452003,"score_gpt":0.39903060898602993,"score_spread":0.3735504838115099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2071311708","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0073933834,0.000074207375,0.98648804,0.00024708043,0.000023979474,0.00015923634,0.00016177312,0.0036658202,0.0017865117],"genre_scores_gemma":[0.041108377,0.00017042113,0.95559144,0.00009644666,0.00001599419,0.0001760302,0.00033802068,0.00133804,0.0011652753],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99741274,0.0008486889,0.0002458637,0.0005576501,0.0007989789,0.0001360492],"domain_scores_gemma":[0.9873691,0.0078084413,0.00081286445,0.0021535167,0.0016554291,0.00020070687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006311107,0.0023133175,0.00076269906,0.0019613665,0.0009967026,0.0033796916,0.0028286425,0.0016854802,0.0067265057],"category_scores_gemma":[0.024057094,0.0014427541,0.001587708,0.00082440436,0.0021472843,0.008284136,0.0027823932,0.0036127763,0.0013382955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034635214,0.00030904656,0.0042974805,0.001705862,0.00014622779,0.001798781,0.015988667,0.03461097,0.056197893,0.43689778,0.018653622,0.42904723],"study_design_scores_gemma":[0.00014165184,0.00026538732,0.0013800089,0.00077527884,0.00019090482,0.002056785,0.002634038,0.37452626,0.10146647,0.31766552,0.19862613,0.00027152564],"about_ca_topic_score_codex":0.0023605747,"about_ca_topic_score_gemma":0.003563259,"teacher_disagreement_score":0.0067265057,"about_ca_system_score_codex":0.0014201582,"about_ca_system_score_gemma":0.0028196576,"threshold_uncertainty_score":0.033376694},"labels":[],"label_agreement":null},{"id":"W2071342967","doi":"10.1145/2245276.2231969","title":"Comparative stability of cloned and non-cloned code","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Cloning (programming); Code (set theory); Stability (learning theory); Computer science; Source code; Software; Source lines of code; Programming language; Set (abstract data type)","score_opus":0.05288241831273374,"score_gpt":0.3155559511082728,"score_spread":0.26267353279553907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2071342967","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98917633,0.0009857722,0.0071948315,0.000050613704,0.000022461792,0.000022493756,0.00046143873,0.00014129501,0.0019447984],"genre_scores_gemma":[0.99577725,0.0001896989,0.0025486082,0.00001900017,0.000014316513,0.000022144326,0.0007187901,0.00007318225,0.000636926],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9970806,0.0007251345,0.00027713086,0.000642006,0.0011148411,0.00016023923],"domain_scores_gemma":[0.95155054,0.020122336,0.008890461,0.0049219,0.013476633,0.0010381875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002841441,0.00023443495,0.00037363992,0.0045779236,0.0004201747,0.001244538,0.00063112495,0.00034441016,0.0018186382],"category_scores_gemma":[0.040238973,0.00018602062,0.00039012847,0.002556469,0.0012400919,0.0020829132,0.0009157316,0.00041722308,0.0004470375],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0069510173,0.00030306377,0.48467654,0.0013661628,0.00087841356,0.0010327431,0.008010024,0.009477829,0.23839563,0.018699098,0.0017190018,0.22849046],"study_design_scores_gemma":[0.000049146827,0.0016167746,0.8695834,0.00012754806,0.00036614802,0.0017951643,0.0017041054,0.021872036,0.0865483,0.010366757,0.005872792,0.00009771851],"about_ca_topic_score_codex":0.0013772211,"about_ca_topic_score_gemma":0.0009868254,"teacher_disagreement_score":0.0045779236,"about_ca_system_score_codex":0.00084426015,"about_ca_system_score_gemma":0.00046596196,"threshold_uncertainty_score":0.015027165},"labels":[],"label_agreement":null},{"id":"W2071779362","doi":"10.1145/505894.505906","title":"Practical data exchange for reverse engineering frameworks","year":2001,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Reverse engineering; Computer science; Software engineering; Software; Systems engineering; Database; Engineering; Operating system","score_opus":0.07369380567982925,"score_gpt":0.33189365732011394,"score_spread":0.2581998516402847,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2071779362","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022316482,0.00023688374,0.9883887,0.0011265797,0.00012942703,0.00032394042,0.0001939187,0.0036146217,0.0037542756],"genre_scores_gemma":[0.04185319,0.000471346,0.94912714,0.0005267044,0.00013613884,0.0005414984,0.00087089307,0.0013712267,0.0051018656],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.97581136,0.009798374,0.003960298,0.0024199504,0.0064914143,0.0015185399],"domain_scores_gemma":[0.959592,0.012991903,0.0018179045,0.019382559,0.005418171,0.00079743954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034717076,0.0015330454,0.0014236646,0.002799718,0.0035777262,0.011702211,0.0052641723,0.0047641424,0.011048627],"category_scores_gemma":[0.053157277,0.0020868587,0.0034000734,0.0038082425,0.004032495,0.023477335,0.010488295,0.006595812,0.0044985204],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017133799,0.00012288938,0.0008671472,0.0004406911,0.00004578987,0.00042609457,0.0012711775,0.0064659175,0.0027487276,0.85197395,0.013849963,0.12161633],"study_design_scores_gemma":[0.00019642414,0.00015396642,0.0002464658,0.0006562578,0.000093136325,0.0010986167,0.00076052605,0.060435746,0.017646529,0.5086236,0.4099462,0.00014268848],"about_ca_topic_score_codex":0.0033113777,"about_ca_topic_score_gemma":0.0029269976,"teacher_disagreement_score":0.034717076,"about_ca_system_score_codex":0.0033244875,"about_ca_system_score_gemma":0.005886108,"threshold_uncertainty_score":0.18360358},"labels":[],"label_agreement":null},{"id":"W2072038212","doi":"10.1049/iet-sen:20089010","title":"Software language engineering","year":2008,"lang":"en","type":"editorial","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Computer science; Programming language; Software engineering; Software construction; Software development; Second-generation programming language; Software system; Comparison of multi-paradigm programming languages; Domain-specific language; Software; Fifth-generation programming language; Programming paradigm","score_opus":0.0076433244501118575,"score_gpt":0.2482138613454336,"score_spread":0.24057053689532176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072038212","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00023378676,0.12377581,0.035482526,0.1683696,0.6285414,0.00011731784,0.0002472405,0.0011214523,0.04211081],"genre_scores_gemma":[0.0075662853,0.1734389,0.023591638,0.09488958,0.6103387,0.00044676772,0.0008012064,0.0017354417,0.087191425],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98921216,0.0035050898,0.0018769779,0.0008345605,0.004246734,0.00032440425],"domain_scores_gemma":[0.96075594,0.02452612,0.0010290653,0.0023304636,0.0101817455,0.0011767234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009422344,0.0021779414,0.0018296413,0.0041793375,0.0019838496,0.0067813615,0.0037710562,0.008131109,0.014908793],"category_scores_gemma":[0.03144192,0.00095553574,0.001835121,0.0021075064,0.006643711,0.009166335,0.003910815,0.017919954,0.016887745],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012569316,0.000014180177,0.000028355746,0.0008438536,0.000038416176,0.00019581403,0.00022216376,0.00020057563,0.00023323587,0.048980016,0.8721027,0.077128135],"study_design_scores_gemma":[0.0000059389986,0.000004983375,0.000017681336,0.00025957965,0.0000056020835,0.0001436428,0.00002673568,0.00007442133,0.000081676386,0.014059421,0.98531264,0.000007521099],"about_ca_topic_score_codex":0.00073964667,"about_ca_topic_score_gemma":0.0009319409,"teacher_disagreement_score":0.014908793,"about_ca_system_score_codex":0.0027528035,"about_ca_system_score_gemma":0.0029396259,"threshold_uncertainty_score":0.0498749},"labels":[],"label_agreement":null},{"id":"W2072143380","doi":"10.1007/s10664-015-9361-0","title":"Evaluating the impact of design pattern and anti-pattern dependencies on changes and faults","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal; Queen's University","funders":"Fonds Québécois de la Recherche sur la Nature et les Technologies; Canada Research Chairs","keywords":"Software design pattern; Structural pattern; Maintainability; Computer science; Design pattern; Flexibility (engineering); Architectural pattern; Data mining; Software design; Software; Software engineering; Software development; Programming language; Mathematics; Statistics","score_opus":0.1456268245412086,"score_gpt":0.381946844692815,"score_spread":0.23632002015160639,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072143380","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99767333,0.00006247698,0.0015344806,0.000043630025,0.000006214417,0.000018390056,0.00015548793,0.000028558909,0.00047738315],"genre_scores_gemma":[0.99821955,0.000026899115,0.0013009548,0.000010177962,0.0000042850515,0.000011875188,0.00023599374,0.000008996217,0.00018128172],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9934384,0.002939383,0.0005505039,0.0010225208,0.0015974367,0.000451806],"domain_scores_gemma":[0.55813384,0.4070131,0.018788682,0.00803318,0.006345857,0.0016853496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007817053,0.00057843403,0.00038214235,0.001579297,0.00029188304,0.0008285827,0.00097081694,0.0009020856,0.0022279676],"category_scores_gemma":[0.119632035,0.00034983057,0.00076677196,0.0011531146,0.00078420324,0.0021359092,0.00049693324,0.0013260002,0.00020548185],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0057651578,0.003569873,0.7863784,0.0002959987,0.00093956356,0.00029416688,0.00028427612,0.10370108,0.010143799,0.0014542062,0.0003659732,0.086807504],"study_design_scores_gemma":[0.00024730724,0.00663626,0.6837173,0.000047581758,0.00087779766,0.00033863456,0.0005827688,0.2877829,0.015787119,0.003251482,0.0006667935,0.00006415717],"about_ca_topic_score_codex":0.0029344524,"about_ca_topic_score_gemma":0.0055516623,"teacher_disagreement_score":0.007817053,"about_ca_system_score_codex":0.0009415587,"about_ca_system_score_gemma":0.0013098345,"threshold_uncertainty_score":0.041341066},"labels":[],"label_agreement":null},{"id":"W2072481344","doi":"10.1080/10798587.2014.960229","title":"Perception-Based Software Release Planning","year":2014,"lang":"en","type":"article","venue":"Intelligent Automation & Soft Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Perception; Software; Software engineering; Human–computer interaction; Programming language; Psychology","score_opus":0.02159196918474365,"score_gpt":0.28509588314987483,"score_spread":0.2635039139651312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072481344","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025181236,0.00036566713,0.9674415,0.0003644554,0.000035990437,0.00012261135,0.00007538062,0.0007686199,0.005644475],"genre_scores_gemma":[0.6755952,0.0003671224,0.3217636,0.00007724159,0.000029136372,0.00013740642,0.00023102289,0.00014323698,0.0016559858],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974567,0.0009488437,0.00012371945,0.00046631737,0.000832393,0.00017216103],"domain_scores_gemma":[0.99568754,0.0022149715,0.00073140673,0.0003744638,0.0007396325,0.00025199336],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035100575,0.0007613559,0.00059197034,0.001250288,0.0005677125,0.0021653126,0.0013904831,0.00061586883,0.0026704294],"category_scores_gemma":[0.010128093,0.00061918295,0.0009153883,0.0007227541,0.000993422,0.0025645006,0.0014630862,0.001254738,0.00048280697],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045171988,0.00034767998,0.011479416,0.00059754314,0.00021739028,0.0004429866,0.003276795,0.42443177,0.01615033,0.0967577,0.004561045,0.44128555],"study_design_scores_gemma":[0.000051054663,0.00024878554,0.0056846444,0.000078771234,0.0000706577,0.00012576568,0.0006379113,0.93608856,0.004668662,0.046266567,0.0059731347,0.0001055037],"about_ca_topic_score_codex":0.007571982,"about_ca_topic_score_gemma":0.006903358,"teacher_disagreement_score":0.007571982,"about_ca_system_score_codex":0.0012964988,"about_ca_system_score_gemma":0.0020395925,"threshold_uncertainty_score":0.018563211},"labels":[],"label_agreement":null},{"id":"W2072632642","doi":"10.1145/512035.512062","title":"Tracking structural evolution using origin analysis","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Software; Software system; Term (time); Software architecture; Software engineering; Work (physics); Architecture; Software construction; Engineering; Programming language; Geography","score_opus":0.054518655594594144,"score_gpt":0.30327786483544283,"score_spread":0.24875920924084868,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072632642","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09550074,0.0004493612,0.89705366,0.00016400442,0.00003663966,0.000050032784,0.00033382865,0.0040786183,0.0023331046],"genre_scores_gemma":[0.61773604,0.0004322701,0.3781377,0.000046809062,0.000028610097,0.000098783414,0.0009057211,0.00054803,0.0020660406],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99887985,0.00023043956,0.000053303153,0.00033015365,0.00042603115,0.000080183505],"domain_scores_gemma":[0.9946642,0.0020279645,0.0010996301,0.0011835117,0.00089842285,0.00012622944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014369149,0.000526707,0.0005613945,0.0033508819,0.0007383691,0.0012943527,0.0012970843,0.0010811945,0.0016100033],"category_scores_gemma":[0.0094099,0.0005072859,0.0006043625,0.0019914224,0.00097484735,0.0024193826,0.002080124,0.0013353735,0.00047117547],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056297006,0.00019935721,0.09957415,0.00048087392,0.00023591662,0.0011159567,0.0033902992,0.19284299,0.07703513,0.100203015,0.0041502146,0.5202092],"study_design_scores_gemma":[0.000029192364,0.00008413155,0.009226039,0.000034192075,0.00007875424,0.00031884763,0.00018242584,0.9228685,0.02241812,0.035676174,0.00903988,0.000043753145],"about_ca_topic_score_codex":0.0033806795,"about_ca_topic_score_gemma":0.0023624047,"teacher_disagreement_score":0.0033806795,"about_ca_system_score_codex":0.0007626329,"about_ca_system_score_gemma":0.0005566937,"threshold_uncertainty_score":0.0075992346},"labels":[],"label_agreement":null},{"id":"W2072825466","doi":"10.1145/1297846.1297930","title":"P&lt;scp&gt;TIDEJ&lt;/scp&gt; and D&lt;scp&gt;ECOR&lt;/scp&gt;","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Identification (biology); Computer science; Suite; Code (set theory); Focus (optics); Computational biology; Programming language; Biology; Physics; Geography","score_opus":0.015030947749156372,"score_gpt":0.2552390446949922,"score_spread":0.24020809694583584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072825466","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0079387445,0.002223117,0.5566358,0.0039459676,0.0019643279,0.0005725612,0.011018694,0.07345809,0.34224278],"genre_scores_gemma":[0.059813533,0.0021609324,0.32072583,0.0009654224,0.00068413216,0.00052864145,0.03223936,0.03646483,0.54641736],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99781704,0.00027128466,0.00009196897,0.0005337097,0.0010897588,0.00019614377],"domain_scores_gemma":[0.99425644,0.001352158,0.0003151115,0.0014754557,0.0018513458,0.00074948726],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025426731,0.0014555057,0.0009541161,0.0023816016,0.0010202728,0.003514919,0.0020725974,0.0018652874,0.22289671],"category_scores_gemma":[0.006200214,0.00092540734,0.0008791793,0.0031653726,0.0009890109,0.0035456342,0.0027825057,0.002271426,0.14415489],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027522465,0.00018498868,0.0010340334,0.00047620953,0.000027757076,0.00031425437,0.00025604267,0.001832051,0.012630762,0.06262688,0.32051253,0.5998292],"study_design_scores_gemma":[0.00009126513,0.00008485454,0.0011657951,0.0000788877,0.000014708823,0.0005494268,0.00005087824,0.01072641,0.016356096,0.0128619615,0.95797133,0.000048443093],"about_ca_topic_score_codex":0.0042413436,"about_ca_topic_score_gemma":0.005212023,"teacher_disagreement_score":0.22289671,"about_ca_system_score_codex":0.0013605704,"about_ca_system_score_gemma":0.001932052,"threshold_uncertainty_score":0.74566376},"labels":[],"label_agreement":null},{"id":"W2073102918","doi":"10.1007/s11219-014-9238-2","title":"Studying the relationship between source code quality and mobile platform dependence","year":2014,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Android (operating system); Source code; Software quality; Computer science; Software quality assurance; Software; Operating system; Source lines of code; Mobile apps; Mobile device; Mobile computing; Quality assurance; Software development; World Wide Web; Engineering","score_opus":0.152923259708518,"score_gpt":0.38289060231190514,"score_spread":0.22996734260338714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073102918","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9986835,0.00009024021,0.0005867184,0.00007360908,0.0000018585881,0.0000032191897,0.000039048475,0.0000092741275,0.0005125532],"genre_scores_gemma":[0.99926704,0.000042560998,0.0002982651,0.000011350882,0.000004290886,0.0000023605492,0.00006449439,0.000009105524,0.0003005766],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99849224,0.00052598654,0.000113224065,0.00020837908,0.00047123723,0.00018899485],"domain_scores_gemma":[0.8319708,0.12849906,0.02551204,0.0035698,0.007920404,0.0025278763],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023792926,0.0002692452,0.00019720233,0.001788352,0.0003628783,0.0010079637,0.0005978967,0.00060429773,0.002727931],"category_scores_gemma":[0.058522403,0.00031867484,0.00045559986,0.0020926571,0.00055852573,0.0019449948,0.0006875952,0.0015391022,0.0003291913],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017810545,0.00017820043,0.9900961,0.00002297915,0.000105799394,0.00006419946,0.00021334228,0.00091970916,0.0013213976,0.00038992046,0.00008833768,0.0064219525],"study_design_scores_gemma":[0.000008326422,0.00020952946,0.9890274,0.000013100275,0.000103263286,0.00012770019,0.00039433673,0.00822966,0.0011543201,0.00052035507,0.00020243901,0.000009542185],"about_ca_topic_score_codex":0.007842065,"about_ca_topic_score_gemma":0.013678543,"teacher_disagreement_score":0.007842065,"about_ca_system_score_codex":0.00067027385,"about_ca_system_score_gemma":0.00085211714,"threshold_uncertainty_score":0.0155928135},"labels":[],"label_agreement":null},{"id":"W2073372788","doi":"10.1109/nafips.2006.365420","title":"Predicting Qualitative Assessments Using Fuzzy Aggregation","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; University of Manitoba; National Research Council Canada","funders":"","keywords":"Maintainability; Software metric; Computer science; Software; Data mining; Machine learning; Software system; Extensibility; Software sizing; Classifier (UML); Fuzzy logic; Software construction; Component-based software engineering; Artificial intelligence; Software engineering; Programming language","score_opus":0.052119528878357954,"score_gpt":0.38612388590540614,"score_spread":0.33400435702704817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073372788","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30103415,0.0002991333,0.69280523,0.00036167007,0.000063873515,0.00017940578,0.00024932445,0.0008459555,0.0041613462],"genre_scores_gemma":[0.8876508,0.000093150506,0.111477874,0.000046939633,0.00003490104,0.00008520836,0.00015857998,0.000019356212,0.00043326325],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99652267,0.0009753322,0.00024763725,0.0005166147,0.0015481919,0.00018964258],"domain_scores_gemma":[0.986265,0.008235005,0.0011794879,0.0008454281,0.0032248558,0.0002501558],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007861763,0.0013229297,0.0010658802,0.0037276559,0.00072562264,0.0020744177,0.00076890487,0.0008376661,0.0011224065],"category_scores_gemma":[0.024276752,0.00033008127,0.0009869237,0.0014292805,0.0006728509,0.002222198,0.0011655578,0.001166354,0.00039399526],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001003866,0.00028186958,0.04163961,0.00021713396,0.0003861643,0.00026122728,0.0010050728,0.50146896,0.0192592,0.00836841,0.0016706297,0.42443782],"study_design_scores_gemma":[0.000015122983,0.00014852056,0.004924882,0.00002975313,0.000062762636,0.000034884328,0.00011755802,0.979471,0.005364212,0.009448351,0.00033729136,0.000045646684],"about_ca_topic_score_codex":0.0048351195,"about_ca_topic_score_gemma":0.0034196915,"teacher_disagreement_score":0.007861763,"about_ca_system_score_codex":0.0011835216,"about_ca_system_score_gemma":0.0007750043,"threshold_uncertainty_score":0.04157746},"labels":[],"label_agreement":null},{"id":"W2073418641","doi":"10.1145/568760.568762","title":"On the many ways software engineering can benefit from knowledge engineering","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software development; Domain knowledge; Social software engineering; Knowledge management; Personal software process; Software engineering; Team software process; Software; Software construction","score_opus":0.025180664493295545,"score_gpt":0.20946584553694744,"score_spread":0.1842851810436519,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073418641","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071182526,0.10150838,0.25126022,0.38584486,0.0030941514,0.00013951655,0.00018017586,0.00064309704,0.25021142],"genre_scores_gemma":[0.3248103,0.19662422,0.32626224,0.0831409,0.013816313,0.0006433484,0.00042892838,0.0009367603,0.05333697],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9819162,0.009702571,0.0008176163,0.0012361159,0.0047677993,0.001559754],"domain_scores_gemma":[0.95939714,0.031680215,0.0009692367,0.003971324,0.0028146403,0.0011674361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02017897,0.0022927553,0.0017799719,0.007745489,0.00655385,0.025663266,0.003209545,0.0112351915,0.011006972],"category_scores_gemma":[0.028036704,0.0015183092,0.001774757,0.008650056,0.03940925,0.064719886,0.016486017,0.012477361,0.004690585],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002282636,0.000037822483,0.00027184287,0.0003433412,0.000027087952,0.000167559,0.0039063403,0.00062427786,0.00014786777,0.9273243,0.01185749,0.05526924],"study_design_scores_gemma":[0.000008773084,0.000018652674,0.00015155142,0.0005955468,0.000007154158,0.00012718145,0.0015405146,0.0004312611,0.00014404235,0.88433665,0.11260342,0.000035269557],"about_ca_topic_score_codex":0.003459052,"about_ca_topic_score_gemma":0.0036031238,"teacher_disagreement_score":0.025663266,"about_ca_system_score_codex":0.0042285346,"about_ca_system_score_gemma":0.003645455,"threshold_uncertainty_score":0.106717885},"labels":[],"label_agreement":null},{"id":"W2073503206","doi":"10.1007/s10796-013-9428-7","title":"Aspectualization of code clones—an algorithmic approach","year":2013,"lang":"en","type":"article","venue":"Information Systems Frontiers","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Maintainability; Computer science; Source code; Reusability; Modularity (biology); Programming language; Code (set theory); KPI-driven code analysis; Software maintenance; Redundant code; Unreachable code; Code refactoring; Aspect-oriented programming; Static program analysis; Code generation; Software engineering; Software; Software development; Operating system; Key (lock); Set (abstract data type)","score_opus":0.01175113816486471,"score_gpt":0.22185914219804184,"score_spread":0.21010800403317714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073503206","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011278887,0.00005994209,0.9856587,0.00015352377,0.000013431885,0.0000903622,0.000027808192,0.0005466792,0.0021706647],"genre_scores_gemma":[0.22187562,0.00022133716,0.7739339,0.00012893639,0.00006316326,0.00018813595,0.00033501428,0.0005441408,0.0027097687],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955094,0.001164717,0.0003424177,0.00085835735,0.0016685225,0.00045662504],"domain_scores_gemma":[0.98697066,0.006373649,0.0008619111,0.004094952,0.0015078171,0.00019106682],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035216978,0.00082736654,0.001131295,0.002606032,0.001567038,0.0043882243,0.003466285,0.0015810024,0.0037236041],"category_scores_gemma":[0.022558369,0.0010360413,0.0024031638,0.0027222808,0.0035336458,0.005901192,0.0040332,0.002517883,0.0006210944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015874006,0.00024065103,0.004401728,0.00031340553,0.0000901169,0.00025684648,0.0010910492,0.05943911,0.010717325,0.5797504,0.0025289075,0.34101176],"study_design_scores_gemma":[0.000043492735,0.00008889372,0.0011257385,0.00011352528,0.00010600668,0.00036886398,0.00038402068,0.43439826,0.013554544,0.54040426,0.009366783,0.000045630055],"about_ca_topic_score_codex":0.0016389895,"about_ca_topic_score_gemma":0.002419476,"teacher_disagreement_score":0.0043882243,"about_ca_system_score_codex":0.0010545377,"about_ca_system_score_gemma":0.002328899,"threshold_uncertainty_score":0.018624723},"labels":[],"label_agreement":null},{"id":"W2073798166","doi":"10.1142/s0218194008003532","title":"SOFTWARE EFFORT ESTIMATION BY ANALOGY USING ATTRIBUTE SELECTION BASED ON ROUGH SET ANALYSIS","year":2008,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Weighting; Data mining; Selection (genetic algorithm); Computer science; Set (abstract data type); Analogy; Estimation; Rough set; Artificial intelligence; Machine learning; Engineering","score_opus":0.017442438192425847,"score_gpt":0.2725042033759681,"score_spread":0.25506176518354223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073798166","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.061000116,0.0001418111,0.93759155,0.00009873403,0.000015783633,0.000116808274,0.000077829485,0.00029214032,0.0006652631],"genre_scores_gemma":[0.563766,0.00015487625,0.4351901,0.000026388523,0.000023863475,0.00027347048,0.00024951875,0.000026755813,0.00028905208],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99388015,0.003218297,0.0004792736,0.00053813716,0.0017266022,0.00015748975],"domain_scores_gemma":[0.9863073,0.0098785665,0.0011305441,0.0009982252,0.0015377393,0.00014758011],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061866483,0.0009701849,0.0020470512,0.0044495543,0.00058455236,0.0016850589,0.001188549,0.0005876936,0.00064064265],"category_scores_gemma":[0.021935202,0.00048169793,0.0017071571,0.002944429,0.0004902846,0.0022921034,0.0012453893,0.0009440653,0.00015635462],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021942852,0.00024875833,0.019306649,0.00028801142,0.00041945156,0.00021136332,0.00053495297,0.629682,0.003786451,0.019279396,0.001352457,0.32467112],"study_design_scores_gemma":[0.000024397094,0.00009284768,0.0025824204,0.000015568683,0.000052197498,0.000063113635,0.000054894615,0.9792859,0.0018105586,0.01544353,0.0005391601,0.000035502235],"about_ca_topic_score_codex":0.002111389,"about_ca_topic_score_gemma":0.0013811778,"teacher_disagreement_score":0.0061866483,"about_ca_system_score_codex":0.0008849225,"about_ca_system_score_gemma":0.0011959694,"threshold_uncertainty_score":0.03271854},"labels":[],"label_agreement":null},{"id":"W2073838820","doi":"10.1109/suite.2012.6225471","title":"Program analysis using interactive and visual querying","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Set (abstract data type); Selection (genetic algorithm); Information retrieval; Interactive visual analysis; Process (computing); Comprehension; Program comprehension; Interactive visualization; Visualization; Visual analytics; Data mining; Human–computer interaction; Programming language; Artificial intelligence; Software; Software system","score_opus":0.0294822188261601,"score_gpt":0.3705284882650008,"score_spread":0.3410462694388407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073838820","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017714775,0.000029641764,0.99299,0.00010649772,0.0000036132417,0.00006982959,0.000037224476,0.004141994,0.00084966765],"genre_scores_gemma":[0.07579224,0.00012865319,0.92080057,0.00017598682,0.00003682922,0.00031899294,0.00037627332,0.0009981462,0.0013723196],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99062437,0.0039819344,0.00062237604,0.001478728,0.0027906261,0.000502022],"domain_scores_gemma":[0.9821698,0.012287698,0.00072338263,0.0029913601,0.0013681047,0.00045969643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077648773,0.002130289,0.0016928302,0.0035831206,0.0010123783,0.0072321203,0.0055689504,0.0020248957,0.0062965816],"category_scores_gemma":[0.01954419,0.0010445009,0.002469385,0.0020482342,0.004646296,0.009651375,0.006158161,0.0021490536,0.0013728191],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011838519,0.00069680926,0.0037263702,0.0012787019,0.00038166452,0.00075327884,0.011916725,0.03665991,0.06456097,0.32559398,0.012697261,0.5405505],"study_design_scores_gemma":[0.00028004384,0.00046833177,0.0014080235,0.00023140555,0.00021477853,0.00083003263,0.001463156,0.55980587,0.049894657,0.33652204,0.048559453,0.000322232],"about_ca_topic_score_codex":0.0035283575,"about_ca_topic_score_gemma":0.0028599983,"teacher_disagreement_score":0.0077648773,"about_ca_system_score_codex":0.0012327385,"about_ca_system_score_gemma":0.0014843359,"threshold_uncertainty_score":0.041065097},"labels":[],"label_agreement":null},{"id":"W2074019008","doi":"10.1145/2597073.2597098","title":"Works for me! characterizing non-reproducible bug reports","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software bug; Root cause; Software regression; Software; Set (abstract data type); Software engineering; Software quality; Software development; Reliability engineering; Programming language; Engineering","score_opus":0.018621522337595973,"score_gpt":0.2678200077974596,"score_spread":0.24919848545986364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2074019008","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9018194,0.0047952295,0.04344129,0.007048874,0.00052961987,0.00062800647,0.017667742,0.0048423987,0.019227581],"genre_scores_gemma":[0.92179257,0.0017967473,0.04865688,0.0012910716,0.00030408692,0.00046667634,0.012659365,0.0012223807,0.011810319],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99457186,0.0014402005,0.0004577231,0.0009009253,0.0023515548,0.00027762112],"domain_scores_gemma":[0.9314417,0.025366554,0.02028502,0.009465554,0.011391307,0.002049799],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052687707,0.0007383462,0.00047439281,0.0063753976,0.0009915135,0.002120631,0.0010937243,0.00094279734,0.0048594647],"category_scores_gemma":[0.051731642,0.00055247184,0.0006125636,0.0047941953,0.0007062837,0.0023379358,0.001349433,0.0006413293,0.0053967936],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035754705,0.00029565324,0.55400455,0.0009520041,0.00031852158,0.0010721559,0.003754512,0.0013161984,0.006561384,0.0022643683,0.059945453,0.36915764],"study_design_scores_gemma":[0.00009695292,0.0008612196,0.84000415,0.0007571472,0.0002817158,0.004617736,0.0047581363,0.011375397,0.008463904,0.007916838,0.12065076,0.00021612324],"about_ca_topic_score_codex":0.0017418202,"about_ca_topic_score_gemma":0.0031079503,"teacher_disagreement_score":0.0063753976,"about_ca_system_score_codex":0.00045518484,"about_ca_system_score_gemma":0.00063200045,"threshold_uncertainty_score":0.027864218},"labels":[],"label_agreement":null},{"id":"W2075099535","doi":"10.1109/scam.2014.29","title":"Supplementary Bug Fixes vs. Re-opened Bugs","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Software bug; Commit; Eclipse; Computer science; Software regression; Security bug; Debugging; Software; Software engineering; Programming language; Software development; Database; Software quality; Operating system","score_opus":0.013495630049182986,"score_gpt":0.2585956877137616,"score_spread":0.24510005766457865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2075099535","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94218504,0.0038224938,0.039806902,0.00037906264,0.00022541244,0.00020553848,0.0058355154,0.0030110334,0.004529003],"genre_scores_gemma":[0.98754406,0.0006217107,0.0072325883,0.00006358889,0.00006973911,0.00005891958,0.0034520347,0.00009535578,0.0008620283],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99532324,0.0006299099,0.0004354417,0.0019524321,0.0013638078,0.00029521625],"domain_scores_gemma":[0.94240254,0.038728252,0.009281865,0.005054352,0.0038531616,0.00067973044],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005000123,0.0012626919,0.0009191377,0.006431603,0.0004778837,0.0022007318,0.0012281314,0.0012750104,0.0018611896],"category_scores_gemma":[0.04457707,0.0006292768,0.0013831662,0.003047377,0.0009721761,0.0038521385,0.0014106252,0.0020035221,0.0008074948],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010170655,0.00037347563,0.73098975,0.0011403108,0.0007046178,0.0008032145,0.0019462585,0.042716086,0.003888326,0.0031844613,0.007398956,0.20583746],"study_design_scores_gemma":[0.00010577407,0.00068804965,0.60641974,0.0006586607,0.00064309675,0.001951242,0.0010326849,0.35769072,0.008410988,0.008838268,0.013340911,0.00021986986],"about_ca_topic_score_codex":0.007611429,"about_ca_topic_score_gemma":0.00857654,"teacher_disagreement_score":0.007611429,"about_ca_system_score_codex":0.0008341264,"about_ca_system_score_gemma":0.0006227345,"threshold_uncertainty_score":0.026443541},"labels":[],"label_agreement":null},{"id":"W2075269190","doi":"10.1007/s10664-014-9350-8","title":"Linguistic antipatterns: what they are and how developers perceive them","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":116,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Lexicon; Documentation; Source code; Cognitive dissonance; Computer science; Open source; Code (set theory); Empirical research; Affect (linguistics); Code review; Data science; Linguistics; Psychology; Artificial intelligence; Software development; Static program analysis; Social psychology; Software; Programming language; Communication; Epistemology","score_opus":0.064961403368231,"score_gpt":0.28485054955810457,"score_spread":0.21988914618987357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2075269190","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8677104,0.0006497867,0.08680447,0.0063673304,0.00014930079,0.00008770213,0.00024423256,0.0010462027,0.036940567],"genre_scores_gemma":[0.9760399,0.00027058745,0.019416343,0.00059609086,0.000054254375,0.00007086951,0.00018143734,0.00048590032,0.0028846038],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9929349,0.003056821,0.0005336074,0.00072047405,0.0024490906,0.00030507345],"domain_scores_gemma":[0.9493569,0.025655616,0.009115425,0.005201783,0.009469068,0.0012010881],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050787334,0.0003302138,0.00033196656,0.001464488,0.0009523791,0.0041667726,0.0007893291,0.0018530937,0.0032090703],"category_scores_gemma":[0.049852934,0.0006791536,0.00021097859,0.0010621698,0.0024716365,0.009136185,0.0020299493,0.0017453216,0.0007937949],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043220524,0.0003173319,0.34849787,0.00089623657,0.000104799845,0.0011336685,0.1929518,0.0007535458,0.064389646,0.089545704,0.010343952,0.29063323],"study_design_scores_gemma":[0.00014651663,0.0005513046,0.45618743,0.0010307007,0.00037786507,0.0036641317,0.19156946,0.022423,0.028994408,0.18630971,0.10848773,0.0002576612],"about_ca_topic_score_codex":0.0020137767,"about_ca_topic_score_gemma":0.0025442168,"teacher_disagreement_score":0.0050787334,"about_ca_system_score_codex":0.00067583774,"about_ca_system_score_gemma":0.0013805954,"threshold_uncertainty_score":0.026859224},"labels":[],"label_agreement":null},{"id":"W2075646905","doi":"10.1007/s00500-004-0442-z","title":"A soft computing framework for software effort estimation","year":2005,"lang":"en","type":"article","venue":"Soft Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Motorola (Canada); Western University","funders":"","keywords":"Computer science; Interpretability; Soft computing; Machine learning; Data mining; Robustness (evolution); Fuzzy logic; Adaptive neuro fuzzy inference system; Software; Artificial intelligence; Inference; Artificial neural network; Fuzzy control system","score_opus":0.01964834430935424,"score_gpt":0.30361178865353744,"score_spread":0.2839634443441832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2075646905","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00063070195,0.0000788665,0.99867624,0.00008211133,0.000013200708,0.000013135707,0.000018422594,0.00010178586,0.00038554127],"genre_scores_gemma":[0.15439233,0.00055794825,0.84021455,0.00019027402,0.00027523938,0.00031288568,0.00020825713,0.00014761192,0.0037008997],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99594194,0.0015439208,0.0002514592,0.0005439911,0.0014874833,0.00023114879],"domain_scores_gemma":[0.9912369,0.0058928253,0.0005244747,0.00080414314,0.0013095023,0.00023208484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046301973,0.0013620581,0.0020375059,0.0031604758,0.0010295911,0.0033348412,0.0027454358,0.0015139558,0.0034437322],"category_scores_gemma":[0.018192409,0.0008040671,0.0017428817,0.0035410663,0.0015284554,0.0033766367,0.003027281,0.0031900848,0.0009968987],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000065253786,0.00013283895,0.00081413775,0.00018478434,0.00012290904,0.000104332525,0.00016491457,0.4128084,0.0018439937,0.35641223,0.003308049,0.22403815],"study_design_scores_gemma":[0.0000065598992,0.00001880896,0.00011720964,0.000020813426,0.000016988444,0.000024891364,0.000014537597,0.85466754,0.0005437137,0.1432956,0.0012593656,0.000013916036],"about_ca_topic_score_codex":0.005008307,"about_ca_topic_score_gemma":0.005148352,"teacher_disagreement_score":0.005008307,"about_ca_system_score_codex":0.0014290594,"about_ca_system_score_gemma":0.0022611937,"threshold_uncertainty_score":0.024487078},"labels":[],"label_agreement":null},{"id":"W2076091918","doi":"10.1016/j.jss.2010.07.032","title":"Adaptive ridge regression system for software cost estimating on multi-collinear datasets","year":2010,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Saskatchewan; University of Wisconsin-Madison","keywords":"Collinearity; Regression; Computer science; Regression analysis; Linear regression; Machine learning; Proper linear model; Data mining; Artificial intelligence; Polynomial regression; Stepwise regression; Robust regression; Local regression; Regression diagnostic; Software; Statistics; Mathematics","score_opus":0.03813444071146832,"score_gpt":0.30601431565287796,"score_spread":0.26787987494140963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2076091918","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037884866,0.000048248345,0.99311817,0.000029805387,0.000013301083,0.00003060698,0.00010028686,0.0027687955,0.000102352074],"genre_scores_gemma":[0.072667785,0.000094487834,0.92394423,0.00004560213,0.00003536245,0.00030575166,0.0007227268,0.00052649266,0.0016575692],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974197,0.0012000174,0.00020347396,0.00048430325,0.000558374,0.000134133],"domain_scores_gemma":[0.9944159,0.0027044003,0.00038542977,0.001182749,0.0012210222,0.00009063416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005311409,0.0007859277,0.0018782368,0.001528061,0.00037702263,0.0009209899,0.0024271137,0.0013091724,0.00411834],"category_scores_gemma":[0.014039819,0.00083111273,0.0012892951,0.0022210036,0.00036106232,0.0012167623,0.001423788,0.0023865304,0.0024823816],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004996435,0.00031482082,0.003621244,0.00018926045,0.00043403322,0.00011464905,0.00009845387,0.16833502,0.0112036895,0.007036763,0.009066373,0.799086],"study_design_scores_gemma":[0.000040491537,0.000045037283,0.00081811467,0.00000655784,0.00002534292,0.000032094707,0.000010036111,0.9935868,0.0018672332,0.0024377573,0.0011147413,0.000015651754],"about_ca_topic_score_codex":0.0037433642,"about_ca_topic_score_gemma":0.006341508,"teacher_disagreement_score":0.005311409,"about_ca_system_score_codex":0.00043132034,"about_ca_system_score_gemma":0.0015031856,"threshold_uncertainty_score":0.028089762},"labels":[],"label_agreement":null},{"id":"W2076141290","doi":"10.1145/1753196.1753199","title":"DEQUALITE","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science","score_opus":0.04199393215960748,"score_gpt":0.2796160460042001,"score_spread":0.23762211384459261,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2076141290","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018675545,0.00069838244,0.81194127,0.0014447137,0.00048318662,0.0009258132,0.013796853,0.077086434,0.07494776],"genre_scores_gemma":[0.12259841,0.0007499599,0.75320876,0.0008854858,0.00010936568,0.0013120797,0.039318416,0.014668274,0.06714915],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9957404,0.0006471874,0.00043206685,0.00077572995,0.002153817,0.00025077493],"domain_scores_gemma":[0.9931114,0.0019151949,0.00061741617,0.002146869,0.0020089091,0.00020024448],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003418629,0.0014627422,0.0008838384,0.002526906,0.0007155276,0.0035618306,0.0028961254,0.0016301707,0.03250028],"category_scores_gemma":[0.019767009,0.0010689537,0.0014778818,0.0019803566,0.0006564986,0.004731662,0.0033374256,0.0024557365,0.014388184],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053857337,0.00037909055,0.0149685675,0.0013369464,0.00016310203,0.00021307582,0.00073675497,0.020893687,0.00505164,0.1261093,0.17476775,0.65484154],"study_design_scores_gemma":[0.00023865116,0.00031472917,0.0047672847,0.00031713926,0.00008145524,0.0008745024,0.00022452773,0.18938863,0.018344607,0.07690358,0.7084181,0.00012673979],"about_ca_topic_score_codex":0.0025022943,"about_ca_topic_score_gemma":0.0045525446,"teacher_disagreement_score":0.96749973,"about_ca_system_score_codex":0.0014003583,"about_ca_system_score_gemma":0.0015255016,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2077174673","doi":"10.1007/s11219-010-9112-9","title":"A comparative study for estimating software development effort intervals","year":2010,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"National Research Council Canada; Türkiye Bilimsel ve Teknolojik Araştırma Kurumu","keywords":"Computer science; Estimation; Software; Interval estimation; Interval (graph theory); Point (geometry); Point estimation; Data mining; Cluster (spacecraft); Machine learning; Statistics; Confidence interval; Engineering; Mathematics","score_opus":0.08125202920087467,"score_gpt":0.3986317212269463,"score_spread":0.31737969202607164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2077174673","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8570223,0.0032463672,0.13231188,0.00007489122,0.000049432612,0.00022718974,0.0005885473,0.00034177015,0.006137593],"genre_scores_gemma":[0.907927,0.0004216958,0.09008731,0.000017055421,0.000021125144,0.000085565734,0.00059534126,0.00006429385,0.00078070967],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9823245,0.011314342,0.0017251077,0.0012390033,0.0030876002,0.00030935995],"domain_scores_gemma":[0.6082958,0.35717964,0.005749344,0.010553656,0.017396199,0.0008253861],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024084039,0.0007910591,0.0011770193,0.008103463,0.0007761088,0.0018429015,0.0023866347,0.001204038,0.003584798],"category_scores_gemma":[0.15775347,0.0005251797,0.0012649032,0.0072573437,0.00047876206,0.0042481576,0.0009976261,0.0008029966,0.00044354115],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012341227,0.001704809,0.2894871,0.0018189024,0.0012222183,0.00043180786,0.0036556164,0.043287847,0.012639932,0.0054777954,0.0009231398,0.6270097],"study_design_scores_gemma":[0.0008760198,0.021891912,0.42607144,0.00044678198,0.0032509011,0.0021174229,0.004476967,0.5037619,0.027077904,0.0052604964,0.004472729,0.00029544966],"about_ca_topic_score_codex":0.006643042,"about_ca_topic_score_gemma":0.0054913466,"teacher_disagreement_score":0.024084039,"about_ca_system_score_codex":0.0015719907,"about_ca_system_score_gemma":0.00069190277,"threshold_uncertainty_score":0.12737012},"labels":[],"label_agreement":null},{"id":"W2077478714","doi":"10.1007/s10270-007-0055-y","title":"A methodology for the selection of requirements engineering techniques","year":2007,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Requirements engineering; Software engineering; Selection (genetic algorithm); Domain (mathematical analysis); Process (computing); Context (archaeology); Requirements analysis; Systems engineering; Software; Management science; Artificial intelligence; Engineering","score_opus":0.10585584479368089,"score_gpt":0.34902250442522964,"score_spread":0.24316665963154876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2077478714","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00068313605,0.00006511055,0.9970228,0.00011317286,0.000018699995,0.00031873945,0.000097663324,0.0006627338,0.0010179173],"genre_scores_gemma":[0.0054228045,0.000070401904,0.99328834,0.00003847015,0.000016829365,0.00031263055,0.00024971215,0.00012166366,0.00047914265],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97873974,0.009217308,0.0026400262,0.0016896789,0.007157874,0.0005553341],"domain_scores_gemma":[0.96613497,0.017916007,0.0017139518,0.0052669444,0.008423476,0.00054459844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016349413,0.0020509826,0.0016366333,0.008530874,0.0020362574,0.0053456766,0.003228724,0.0018753705,0.008033267],"category_scores_gemma":[0.044836264,0.0017692881,0.0041500493,0.0049214805,0.0015496617,0.0042483453,0.0028735923,0.0037124418,0.0042454842],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000107006446,0.00040754615,0.0016077923,0.0011218278,0.00025524767,0.00045121572,0.0011239357,0.012657502,0.015566973,0.22247139,0.011541586,0.7326879],"study_design_scores_gemma":[0.00043146728,0.0008390333,0.0034676483,0.0018114384,0.00071664376,0.003289188,0.0013019472,0.37257257,0.041408636,0.37310755,0.20055276,0.00050106336],"about_ca_topic_score_codex":0.0020560806,"about_ca_topic_score_gemma":0.003314818,"teacher_disagreement_score":0.016349413,"about_ca_system_score_codex":0.001580213,"about_ca_system_score_gemma":0.004608485,"threshold_uncertainty_score":0.086465},"labels":[],"label_agreement":null},{"id":"W2077616397","doi":"","title":"Modification and Developer Metrics at the Function Level: Metrics for the Study of the Evolution of a Software Project","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software evolution; Computer science; Function point; Software; Software maintenance; Source lines of code; Software engineering; Software metric; Function (biology); Code (set theory); Software project management; Software development; Software construction; Programming language; Set (abstract data type)","score_opus":0.1199557654248466,"score_gpt":0.3181152813373423,"score_spread":0.1981595159124957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2077616397","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.704486,0.0025847007,0.28603688,0.0009042324,0.000050601637,0.00020560379,0.0012037526,0.0007161319,0.0038120174],"genre_scores_gemma":[0.9162354,0.00029081095,0.08232378,0.000024402048,0.000031289193,0.00016139747,0.000606242,0.00007378998,0.00025291793],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.987986,0.006217825,0.0013745484,0.000803303,0.0034076476,0.00021063747],"domain_scores_gemma":[0.90737075,0.0654945,0.014144606,0.0052560577,0.006320028,0.0014140991],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008053359,0.0008298623,0.0007584991,0.012534115,0.0005993953,0.0025712019,0.0008630394,0.0010839351,0.0006147789],"category_scores_gemma":[0.06905668,0.00024897626,0.0006203425,0.014380714,0.0010104972,0.003423703,0.0012945031,0.0010803599,0.0002051033],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014100249,0.00023215453,0.7255765,0.00054174446,0.0003985663,0.00015866176,0.001201202,0.020621771,0.006989935,0.008097122,0.001296076,0.23474532],"study_design_scores_gemma":[0.000033892735,0.0010022945,0.6814999,0.00036609615,0.0002865453,0.0011954632,0.0011658419,0.27479795,0.0107128145,0.021955404,0.006833644,0.00015021741],"about_ca_topic_score_codex":0.0015442175,"about_ca_topic_score_gemma":0.0014307636,"teacher_disagreement_score":0.012534115,"about_ca_system_score_codex":0.0010017648,"about_ca_system_score_gemma":0.0010125948,"threshold_uncertainty_score":0.042590737},"labels":[],"label_agreement":null},{"id":"W2078037959","doi":"10.1109/fie.2012.6462504","title":"// TODO: Help students improve commenting practices","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Documentation; Best practice; Grading (engineering); Source code; Code review; Java; Code (set theory); Set (abstract data type); Internal documentation; Open source; Static program analysis; Software engineering; World Wide Web; Software; Programming language; Software development; Engineering; Software construction","score_opus":0.04291246178722468,"score_gpt":0.3659871521352698,"score_spread":0.3230746903480451,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2078037959","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36519232,0.00062376924,0.2656638,0.007280861,0.0014763469,0.006146306,0.014696436,0.2806961,0.05822403],"genre_scores_gemma":[0.41132113,0.00050567335,0.50050455,0.0019979621,0.00043068238,0.0047546653,0.011008971,0.0135611985,0.055915188],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9947753,0.0020484754,0.00055372296,0.0010321921,0.0012352822,0.0003549983],"domain_scores_gemma":[0.9297943,0.029443074,0.005372283,0.013985393,0.016686985,0.0047179796],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.007325486,0.0017622807,0.001141109,0.0024145779,0.0009377014,0.0028479807,0.0027577195,0.0014771938,0.030179285],"category_scores_gemma":[0.06727658,0.0006853793,0.0006872762,0.001208506,0.00053730374,0.004750529,0.0036160562,0.0016148408,0.021090802],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010648349,0.002880601,0.03970851,0.0015160488,0.00006344134,0.0004533643,0.012606787,0.0008572435,0.026023237,0.0011044277,0.21070306,0.7030184],"study_design_scores_gemma":[0.001124646,0.0042470098,0.09809773,0.0012953142,0.00039799203,0.0017430459,0.020560088,0.05392068,0.1256135,0.010415515,0.68155533,0.0010292894],"about_ca_topic_score_codex":0.00085835165,"about_ca_topic_score_gemma":0.002458455,"teacher_disagreement_score":0.96982074,"about_ca_system_score_codex":0.00067150436,"about_ca_system_score_gemma":0.0013287485,"threshold_uncertainty_score":0.10095978},"labels":[],"label_agreement":null},{"id":"W2078095173","doi":"10.1016/s0164-1212(00)00121-7","title":"An industrial study of reuse, quality, and productivity","year":2001,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":71,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Productivity; Quality (philosophy); Computer science; Interface (matter); Engineering; Operating system; Waste management","score_opus":0.09495754713181716,"score_gpt":0.33974108319004304,"score_spread":0.2447835360582259,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2078095173","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9757951,0.0013681112,0.006534866,0.0005769551,0.000014121214,0.00008448442,0.00005874924,0.000033830707,0.015533856],"genre_scores_gemma":[0.99522865,0.00038471466,0.003517231,0.00003737107,0.00001374038,0.000020953838,0.00003067711,0.000010430525,0.00075625884],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99471855,0.0032397083,0.00023507122,0.00035652617,0.0012701214,0.0001800577],"domain_scores_gemma":[0.88208663,0.09984334,0.0047468105,0.005385365,0.0067518046,0.001186052],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072013894,0.0002751256,0.00025678918,0.0034342522,0.001271832,0.0015425385,0.0009867639,0.00062560063,0.0021022903],"category_scores_gemma":[0.04756309,0.00035182276,0.0003725622,0.0033800756,0.0022940899,0.0025202206,0.001230918,0.0010499515,0.00022786137],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016076589,0.0080961725,0.5035493,0.00061326363,0.0003438908,0.0007901226,0.018256348,0.013772527,0.0096598985,0.0986837,0.0031659324,0.34146118],"study_design_scores_gemma":[0.0006525692,0.011990692,0.7269471,0.00067316135,0.0008681929,0.0014240037,0.019093305,0.09651895,0.022965081,0.08369485,0.035005428,0.00016671576],"about_ca_topic_score_codex":0.009121051,"about_ca_topic_score_gemma":0.009137489,"teacher_disagreement_score":0.009121051,"about_ca_system_score_codex":0.0029087535,"about_ca_system_score_gemma":0.001880302,"threshold_uncertainty_score":0.038085043},"labels":[],"label_agreement":null},{"id":"W2078387967","doi":"10.1109/icsm.2007.4362665","title":"MythSE - myths in software engineering half day ICSM 2007 working session","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Mythology; Session (web analytics); Computer science; Set (abstract data type); Software; Code (set theory); Software engineering; World Wide Web; Multimedia; History; Programming language; Classics","score_opus":0.01572684965865657,"score_gpt":0.26092765371105925,"score_spread":0.2452008040524027,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2078387967","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03368521,0.024979588,0.0136009855,0.6229798,0.13137707,0.0005154515,0.0009374048,0.0009871072,0.1709374],"genre_scores_gemma":[0.20946056,0.02685089,0.016488552,0.1113559,0.05834838,0.0016682968,0.002604891,0.001696336,0.5715261],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976158,0.0009407555,0.00009075263,0.00026573695,0.00075980375,0.00032721134],"domain_scores_gemma":[0.9934223,0.0022676727,0.00022158984,0.00039810984,0.001788889,0.0019013989],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070496886,0.0008380097,0.0007552764,0.0011348687,0.0071575614,0.007034665,0.0010744014,0.0042651836,0.021913346],"category_scores_gemma":[0.008017692,0.00040140047,0.00089049013,0.0006850378,0.0033069488,0.005210594,0.004846601,0.009410758,0.0074368576],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000051980103,0.00019474843,0.0004267133,0.00010819284,0.0000103049515,0.0001398446,0.0057338304,0.00016720599,0.00067215605,0.014529721,0.9380106,0.039954707],"study_design_scores_gemma":[0.000009777762,0.00007176424,0.000729731,0.00025923862,0.0000040163113,0.00012843907,0.005926247,0.00018651258,0.00034907885,0.004845371,0.9874702,0.00001952348],"about_ca_topic_score_codex":0.001377396,"about_ca_topic_score_gemma":0.002567962,"teacher_disagreement_score":0.021913346,"about_ca_system_score_codex":0.0034016052,"about_ca_system_score_gemma":0.0030429126,"threshold_uncertainty_score":0.073307455},"labels":[],"label_agreement":null},{"id":"W2078392793","doi":"10.4018/jssci.2009040104","title":"A Theory of Program Comprehension","year":2009,"lang":"en","type":"article","venue":"International Journal of Software Science and Computational Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"","keywords":"Comprehension; Program comprehension; Computer science; Software; Process (computing); Vision science; Perspective (graphical); Cognitive science; Human–computer interaction; Artificial intelligence; Software system; Programming language; Psychology","score_opus":0.029852218194042573,"score_gpt":0.3386157124770408,"score_spread":0.3087634942829982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2078392793","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00961431,0.0033236465,0.83354205,0.010754243,0.0003371466,0.00027449717,0.0005993694,0.0014453084,0.14010943],"genre_scores_gemma":[0.5601639,0.0052124225,0.3770687,0.006238263,0.0014711445,0.0016292651,0.0022830046,0.0011160823,0.044817124],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.995911,0.001748268,0.00026654088,0.00087597163,0.00086776336,0.00033047615],"domain_scores_gemma":[0.98933136,0.007644726,0.00045194494,0.0010164817,0.0013179438,0.00023758365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004125737,0.0012610807,0.00083818944,0.0027851625,0.0021858248,0.005075065,0.0024586455,0.003490826,0.017695216],"category_scores_gemma":[0.015498008,0.00073179865,0.002800019,0.001915621,0.010041949,0.01821342,0.0032132755,0.0048551937,0.0036129255],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015383881,0.000026188487,0.00021238292,0.00015574391,0.000014359649,0.00007068788,0.0013821069,0.0011790221,0.00029480658,0.9764712,0.0032601722,0.016918084],"study_design_scores_gemma":[0.000018698345,0.00002168984,0.00015172905,0.000079123085,0.000014226199,0.00009568253,0.00019522251,0.0048871306,0.000361437,0.96999145,0.024171414,0.000012291394],"about_ca_topic_score_codex":0.0030982252,"about_ca_topic_score_gemma":0.0013127876,"teacher_disagreement_score":0.017695216,"about_ca_system_score_codex":0.0028443525,"about_ca_system_score_gemma":0.0026830123,"threshold_uncertainty_score":0.059196413},"labels":[],"label_agreement":null},{"id":"W2078515278","doi":"10.5555/2820282.2820288","title":"Detection of software evolution phases based on development activities","year":2015,"lang":"en","type":"article","venue":"International Conference on Program Comprehension","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Software evolution; Computer science; Software development; Granularity; Software; Software construction; Commit; Software engineering; Software sizing; Software analytics; Software metric; Software maintenance; Data mining; Database; Programming language","score_opus":0.08285909543163568,"score_gpt":0.3309184893314647,"score_spread":0.24805939389982906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2078515278","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6189492,0.00089691556,0.37080124,0.00022513345,0.00003706729,0.0005435966,0.0016885767,0.0035625412,0.0032957627],"genre_scores_gemma":[0.72593105,0.00035804076,0.26909444,0.000036868994,0.00002263745,0.00023564122,0.0029958952,0.0002805304,0.001044883],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9977502,0.00044446503,0.00026420795,0.0005105524,0.00088563544,0.00014495829],"domain_scores_gemma":[0.9788681,0.011149164,0.003911883,0.00166939,0.0038041237,0.00059733086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020985054,0.0006151977,0.00053427665,0.008579115,0.00047907548,0.001268542,0.00065155106,0.00052853406,0.00090947404],"category_scores_gemma":[0.017677564,0.00045949247,0.0006180563,0.0036030465,0.00038717943,0.0013358791,0.0009575192,0.00083001226,0.00028829568],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006612702,0.0003553368,0.26585206,0.0009059649,0.00016416657,0.00071253476,0.0032938917,0.026190275,0.057861183,0.006431076,0.002771178,0.6348011],"study_design_scores_gemma":[0.000117693606,0.00075280166,0.3430374,0.00036804407,0.00032759894,0.0018037112,0.0020270825,0.54301816,0.078317426,0.012709512,0.017322069,0.00019854063],"about_ca_topic_score_codex":0.0041725165,"about_ca_topic_score_gemma":0.005663057,"teacher_disagreement_score":0.008579115,"about_ca_system_score_codex":0.0004997996,"about_ca_system_score_gemma":0.001279079,"threshold_uncertainty_score":0.011098087},"labels":[],"label_agreement":null},{"id":"W2079194580","doi":"10.1109/saner.2015.7081815","title":"Measuring the quality of design pattern detection results","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Design pattern; Software design pattern; Pattern detection; Engineering design process; Structural pattern; Quality (philosophy); Task (project management); Process (computing); Specification pattern; Artificial intelligence; Data mining; Software; Software design; Software engineering; Software development; Engineering; Systems engineering; Programming language","score_opus":0.2239854316130997,"score_gpt":0.3308078191398476,"score_spread":0.10682238752674789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079194580","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44688654,0.0033353223,0.5253568,0.001215354,0.00042136796,0.0008881177,0.002234449,0.0120671,0.0075949994],"genre_scores_gemma":[0.6935195,0.00055405166,0.2990493,0.0002552869,0.000107719454,0.00038284823,0.002908266,0.0018805042,0.001342569],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8729921,0.04562463,0.02060947,0.009653957,0.048754565,0.0023653074],"domain_scores_gemma":[0.4101267,0.41253155,0.044347066,0.051292073,0.07872447,0.0029780974],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.080556825,0.0024671168,0.0021226802,0.013063652,0.0011671653,0.0076185632,0.0026792218,0.002907111,0.001919167],"category_scores_gemma":[0.39142928,0.0007422136,0.0017918922,0.004599261,0.0017395414,0.0054302155,0.0036214774,0.001736269,0.0010015768],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035958507,0.0012156516,0.11832488,0.003955209,0.0017004937,0.00085841806,0.0059111156,0.03838752,0.0452223,0.006353015,0.0072898,0.7671858],"study_design_scores_gemma":[0.0010368485,0.0060905246,0.17311355,0.0020549826,0.0027003584,0.0033153787,0.0058483058,0.4821845,0.26655522,0.02940602,0.02645014,0.0012441048],"about_ca_topic_score_codex":0.0016395452,"about_ca_topic_score_gemma":0.0016555698,"teacher_disagreement_score":0.9194432,"about_ca_system_score_codex":0.0014929496,"about_ca_system_score_gemma":0.0017586732,"threshold_uncertainty_score":0.42603028},"labels":[],"label_agreement":null},{"id":"W2079317829","doi":"10.1145/1134285.1134336","title":"Who should fix this bug?","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":921,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Open source; Computer science; Eclipse; Software engineering; Debugging; Software bug; Process (computing); Open-source software development; World Wide Web; Classifier (UML); Usability; Data science; Artificial intelligence; Programming language; Software; Operating system","score_opus":0.023260665400049233,"score_gpt":0.2700901157109861,"score_spread":0.2468294503109369,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079317829","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2617504,0.019042732,0.18935269,0.32815245,0.014453468,0.00090499735,0.0038510188,0.017892297,0.16459998],"genre_scores_gemma":[0.7422233,0.007598136,0.11899375,0.023629941,0.002267473,0.00025995463,0.0026331055,0.0021380084,0.1002562],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9951892,0.001253752,0.0004145192,0.0008821541,0.0017666107,0.00049374456],"domain_scores_gemma":[0.9756217,0.0070719332,0.0050967257,0.0025275175,0.0078779515,0.0018041823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076480405,0.0008070601,0.00084968633,0.0033434946,0.0021480313,0.002790713,0.0012520791,0.0030979733,0.01278852],"category_scores_gemma":[0.063510336,0.0004929235,0.0006972698,0.00182108,0.0013421343,0.0050987587,0.001438488,0.0022525482,0.008499887],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019015922,0.00025024224,0.11439913,0.00046702917,0.00010543949,0.0020239826,0.004755269,0.0006188642,0.0035318448,0.012628457,0.18555743,0.67547214],"study_design_scores_gemma":[0.00011339555,0.0006190806,0.11106926,0.0018514742,0.00035039423,0.0151501885,0.01883375,0.00832997,0.013315227,0.037413687,0.79259795,0.00035563152],"about_ca_topic_score_codex":0.00973923,"about_ca_topic_score_gemma":0.011616897,"teacher_disagreement_score":0.01278852,"about_ca_system_score_codex":0.0016775277,"about_ca_system_score_gemma":0.0038762395,"threshold_uncertainty_score":0.04278189},"labels":[],"label_agreement":null},{"id":"W2079524575","doi":"10.1109/raise.2013.6615204","title":"Handling missing attributes using matrix factorization","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Missing data; Matrix decomposition; Recommender system; Domain (mathematical analysis); Software; Data mining; Machine learning; Domain knowledge; Artificial intelligence; Data science","score_opus":0.039825194383794164,"score_gpt":0.29690347063101274,"score_spread":0.2570782762472186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079524575","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037566973,0.00021316065,0.9950552,0.00019803832,0.000047601552,0.000044552486,0.00013530695,0.00024019931,0.00030924345],"genre_scores_gemma":[0.17523314,0.0006881717,0.81959033,0.00029362633,0.00026102047,0.00034266512,0.0015810787,0.00009385748,0.001916024],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954268,0.0018844309,0.00027380174,0.0010292226,0.0010262226,0.00035961857],"domain_scores_gemma":[0.98564583,0.00955182,0.0011172553,0.0014856408,0.001901059,0.00029835736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049730283,0.0014367432,0.0026045265,0.0019354131,0.0015930071,0.0018307101,0.0021895564,0.0019178733,0.0030172793],"category_scores_gemma":[0.017675087,0.00085873576,0.0020642674,0.0030508996,0.0010344249,0.0037927022,0.0020062402,0.00335157,0.0012648321],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039987513,0.00044719147,0.005075408,0.0007219913,0.00043979133,0.00067804084,0.00093161647,0.38968995,0.005111784,0.047964785,0.018635629,0.52990395],"study_design_scores_gemma":[0.000034261295,0.000103345476,0.00048721136,0.000050453527,0.000047253725,0.0002139949,0.0001722984,0.93572927,0.0013699776,0.058070097,0.0036802571,0.000041534924],"about_ca_topic_score_codex":0.009881017,"about_ca_topic_score_gemma":0.0078623695,"teacher_disagreement_score":0.009881017,"about_ca_system_score_codex":0.00086824986,"about_ca_system_score_gemma":0.0022109004,"threshold_uncertainty_score":0.026300192},"labels":[],"label_agreement":null},{"id":"W2079972013","doi":"10.5555/2664446.2664462","title":"Mining challenge 2012: the Android platform","year":2012,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Android (operating system); Computer science; Data science; World Wide Web; Android application; Software; Android app; Operating system","score_opus":0.027878367104243282,"score_gpt":0.26460729060075805,"score_spread":0.23672892349651475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079972013","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29234445,0.022963798,0.062254928,0.12765819,0.015759427,0.0038936872,0.32046354,0.06511669,0.08954533],"genre_scores_gemma":[0.28544265,0.0053572264,0.089229226,0.008655172,0.003795067,0.0028692782,0.53364074,0.006912999,0.06409771],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9833149,0.0034000212,0.0012161919,0.0023212868,0.008621025,0.0011265654],"domain_scores_gemma":[0.9651196,0.011080197,0.0019294337,0.007895457,0.010706387,0.0032689478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012033563,0.0016545659,0.0018906152,0.0048384825,0.0030983281,0.005481907,0.0031493232,0.0042482405,0.0050166473],"category_scores_gemma":[0.043675374,0.0009332369,0.0016000643,0.0046825847,0.0012570615,0.0060783573,0.006122851,0.0042640655,0.007082395],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059351855,0.00041079224,0.008381545,0.00088787073,0.00013144367,0.0005510005,0.00078151777,0.001460745,0.00232036,0.005187357,0.8924989,0.08679497],"study_design_scores_gemma":[0.00043071044,0.0005357007,0.039245863,0.00042478347,0.00011548254,0.0012236774,0.0016170483,0.028261837,0.008735622,0.009341521,0.9098697,0.00019816846],"about_ca_topic_score_codex":0.020626204,"about_ca_topic_score_gemma":0.02885329,"teacher_disagreement_score":0.020626204,"about_ca_system_score_codex":0.0018791122,"about_ca_system_score_gemma":0.0053015705,"threshold_uncertainty_score":0.063640356},"labels":[],"label_agreement":null},{"id":"W2080002385","doi":"10.1145/1370143.1370148","title":"The benefits and challenges of executable acceptance testing","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Executable; Computer science; Acceptance testing; Software engineering; Programming language","score_opus":0.08238491627779673,"score_gpt":0.2579861573770468,"score_spread":0.17560124109925004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080002385","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07572558,0.009924471,0.8233349,0.039580032,0.00078004674,0.00022628032,0.000080063706,0.0014527574,0.04889593],"genre_scores_gemma":[0.7884979,0.0058864746,0.19330059,0.0049133017,0.0008353555,0.00040803358,0.00015867765,0.0005545064,0.0054451325],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.93548465,0.037038345,0.0023013917,0.0019673351,0.021858845,0.0013494128],"domain_scores_gemma":[0.6742562,0.27337727,0.008305633,0.018642066,0.023657354,0.0017614479],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035392202,0.0013284462,0.0007850515,0.0018382821,0.0012613642,0.004455326,0.0027503502,0.004304014,0.0025485428],"category_scores_gemma":[0.14857751,0.0007647947,0.0008451685,0.0013025253,0.007053679,0.0104155755,0.003669279,0.007011736,0.00088989176],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003051617,0.0002755675,0.012772722,0.0009866597,0.000080915845,0.0017861142,0.0036343387,0.021322997,0.0055320314,0.5401603,0.0048155,0.4083276],"study_design_scores_gemma":[0.00012315447,0.0007007892,0.00677162,0.0013272928,0.00011751145,0.0042248326,0.0026802055,0.0792678,0.007382028,0.8436152,0.053602714,0.00018697827],"about_ca_topic_score_codex":0.0015275999,"about_ca_topic_score_gemma":0.0012211319,"teacher_disagreement_score":0.035392202,"about_ca_system_score_codex":0.0009876479,"about_ca_system_score_gemma":0.0021277484,"threshold_uncertainty_score":0.18717414},"labels":[],"label_agreement":null},{"id":"W2080139099","doi":"10.2308/isys-50809","title":"Business Modeling to Improve Auditor Risk Assessment: An Investigation of Alternative Representations","year":2014,"lang":"en","type":"article","venue":"Journal of Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge; University of Waterloo","funders":"","keywords":"Diagrammatic reasoning; Audit; Presentation (obstetrics); Financial statement; Structuring; Computer science; Representation (politics); Accounting; Statement (logic); Knowledge management; Business; Finance; Linguistics","score_opus":0.01932093057412574,"score_gpt":0.3000902145339554,"score_spread":0.28076928395982964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080139099","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64292526,0.00081426144,0.33954182,0.001732543,0.00011603784,0.0007661157,0.00029560472,0.0016229749,0.012185374],"genre_scores_gemma":[0.7843939,0.00035274684,0.21386142,0.00015393198,0.000021700716,0.00021948782,0.00013530599,0.00011298997,0.00074852013],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97370005,0.022290116,0.0010101628,0.0007496976,0.0019904824,0.00025955672],"domain_scores_gemma":[0.7429626,0.21946342,0.014886224,0.0142511735,0.0075607295,0.0008758551],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021427121,0.0007639876,0.0005245561,0.0016167869,0.00051431573,0.004834553,0.001170772,0.0009343718,0.0029860954],"category_scores_gemma":[0.14804636,0.00049148354,0.00087637495,0.0015891638,0.00092119706,0.0048792968,0.0023835725,0.001474836,0.0003840665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005608021,0.003089353,0.03903841,0.0036433097,0.0004700148,0.00055289984,0.03846075,0.07097177,0.025259042,0.054833204,0.0031132165,0.75496],"study_design_scores_gemma":[0.0013051978,0.0074129626,0.030550156,0.0049614115,0.0012521435,0.0013172555,0.021441288,0.7924573,0.03862954,0.047789976,0.05231178,0.0005710084],"about_ca_topic_score_codex":0.0011619637,"about_ca_topic_score_gemma":0.0011217025,"teacher_disagreement_score":0.021427121,"about_ca_system_score_codex":0.0014659753,"about_ca_system_score_gemma":0.0014315655,"threshold_uncertainty_score":0.1133188},"labels":[],"label_agreement":null},{"id":"W2080382377","doi":"10.1109/icebe.2007.137","title":"Systematic Security Analysis for Service-Oriented Software Architectures","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software security assurance; Software engineering; Security service; Risk analysis (engineering); Security engineering; Computer security; Threat model; Software development; Software architecture; Software; Information security","score_opus":0.011422721304895253,"score_gpt":0.26993344963032617,"score_spread":0.2585107283254309,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080382377","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031785123,0.00032614745,0.96568954,0.00022827095,0.000010945152,0.00019228557,0.000055321656,0.00031531777,0.0013970122],"genre_scores_gemma":[0.4074877,0.0008032611,0.588949,0.00009890976,0.000024647745,0.0006576798,0.0003346099,0.00015472231,0.0014894783],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9950997,0.0022067975,0.0002776779,0.00021644124,0.0020085792,0.0001908354],"domain_scores_gemma":[0.9917235,0.0048715924,0.0007861118,0.0014942057,0.0010464836,0.00007819452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00538967,0.0011384626,0.00076123513,0.0036551198,0.0012712732,0.001623528,0.00072373793,0.0012608616,0.0010616614],"category_scores_gemma":[0.012030973,0.0010016933,0.0027172056,0.0011174871,0.002478908,0.0023960036,0.0016347504,0.0015884694,0.00023394036],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001327617,0.0002655113,0.010005658,0.000860531,0.00023136241,0.0008121512,0.0021324912,0.2704112,0.020498235,0.56105065,0.0016330336,0.13196637],"study_design_scores_gemma":[0.00003834296,0.000148295,0.0013941259,0.00026581148,0.00015035474,0.00024024534,0.0003465391,0.6707987,0.011180626,0.3088346,0.0065565123,0.000045801673],"about_ca_topic_score_codex":0.0031449697,"about_ca_topic_score_gemma":0.0047803214,"teacher_disagreement_score":0.00538967,"about_ca_system_score_codex":0.0013525139,"about_ca_system_score_gemma":0.004210469,"threshold_uncertainty_score":0.028503656},"labels":[],"label_agreement":null},{"id":"W2080505681","doi":"10.1016/j.infsof.2004.08.005","title":"Replicating software engineering experiments: a poisoned chalice or the Holy Grail","year":2004,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Holy Grail; Software; Engineering; Computer science; Software engineering; World Wide Web; Programming language","score_opus":0.0136977408834442,"score_gpt":0.25503565745400436,"score_spread":0.24133791657056017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080505681","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2994825,0.013311587,0.5141412,0.079735234,0.014725678,0.002998211,0.0019431112,0.008039062,0.06562336],"genre_scores_gemma":[0.83772343,0.0025446066,0.11932223,0.018649574,0.00222378,0.002252849,0.0006175344,0.0017075861,0.014958351],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96320474,0.027277129,0.00089878927,0.0030895458,0.0049833097,0.00054639793],"domain_scores_gemma":[0.7263274,0.1126637,0.008899547,0.14140287,0.007861459,0.0028449695],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06323672,0.001442554,0.002498271,0.0014596538,0.0015124083,0.0051893657,0.0050569843,0.005020223,0.0065483013],"category_scores_gemma":[0.25099266,0.0009783179,0.0010756601,0.0010901449,0.011053646,0.010657015,0.005380884,0.0070503335,0.002960169],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009841391,0.0031740374,0.027086781,0.0020801458,0.0027824026,0.0011882882,0.013812174,0.027855946,0.025146684,0.36579132,0.0900256,0.43121526],"study_design_scores_gemma":[0.0024513002,0.004757293,0.008333061,0.0007904514,0.00055716734,0.00050207245,0.0022166423,0.0360754,0.014867676,0.8349151,0.09408878,0.00044515316],"about_ca_topic_score_codex":0.0017030265,"about_ca_topic_score_gemma":0.0011735421,"teacher_disagreement_score":0.9367633,"about_ca_system_score_codex":0.0018480425,"about_ca_system_score_gemma":0.0025010805,"threshold_uncertainty_score":0.33443177},"labels":[],"label_agreement":null},{"id":"W2080534028","doi":"10.1145/1181775.1181779","title":"Questions programmers ask during software evolution tasks","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":259,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Programmer; Categorization; Task (project management); Context (archaeology); Ask price; Program comprehension; Human–computer interaction; Focus (optics); Code (set theory); Software; Data science; Software engineering; Programming language; Software system; Artificial intelligence","score_opus":0.008389132803656189,"score_gpt":0.24183370224981673,"score_spread":0.23344456944616054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080534028","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9672258,0.0005795506,0.02176919,0.0030116593,0.000053157928,0.0002218954,0.00017052116,0.00041266222,0.0065556704],"genre_scores_gemma":[0.9831199,0.0005236901,0.012264194,0.0012682411,0.000051156097,0.0002744585,0.00021023804,0.00012396547,0.0021642006],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.98386616,0.011535773,0.0008348891,0.0010929629,0.0013423066,0.001327972],"domain_scores_gemma":[0.8516256,0.12972553,0.008048789,0.0025365443,0.0050619203,0.0030016236],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012745805,0.0008824803,0.00069784815,0.0015203311,0.002706377,0.0027669535,0.0012729998,0.0049134353,0.0031438011],"category_scores_gemma":[0.10136327,0.0008862882,0.0006221481,0.0009693312,0.0022683602,0.0055354508,0.0032492403,0.0024610688,0.0008459605],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048288418,0.00026050542,0.07568966,0.0006923306,0.00004273702,0.0011433563,0.8600051,0.0005940327,0.012407732,0.0023954273,0.0041591306,0.042127125],"study_design_scores_gemma":[0.0002031724,0.0015184972,0.084533,0.00083336857,0.000113871814,0.0027323163,0.8105723,0.0062696673,0.0087317005,0.008353037,0.075852424,0.00028664345],"about_ca_topic_score_codex":0.001967105,"about_ca_topic_score_gemma":0.0016366794,"teacher_disagreement_score":0.012745805,"about_ca_system_score_codex":0.0010685694,"about_ca_system_score_gemma":0.0012130857,"threshold_uncertainty_score":0.06740713},"labels":[],"label_agreement":null},{"id":"W2080567578","doi":"10.1109/ccece.2008.4564624","title":"Customizing the capture of software architectural design decisions","year":2008,"lang":"en","type":"article","venue":"Conference proceedings - Canadian Conference on Electrical and Computer Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Software engineering; Documentation; Software; Software development; Systems engineering; Engineering","score_opus":0.035188541321035695,"score_gpt":0.21659986276410367,"score_spread":0.18141132144306799,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080567578","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.050036106,0.00018431374,0.9343516,0.00029182184,0.000037416514,0.001561242,0.00037161875,0.0067319707,0.0064339577],"genre_scores_gemma":[0.1440958,0.000306083,0.8482841,0.0001530096,0.000026944466,0.0013822957,0.0015883324,0.00071640854,0.003446984],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9851151,0.006204001,0.001960805,0.0017832412,0.004198722,0.0007380327],"domain_scores_gemma":[0.90018743,0.045652915,0.0055584316,0.04016813,0.0076013235,0.00083173404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015956651,0.001420355,0.00089800014,0.0051533757,0.0007860138,0.0046397424,0.0024088065,0.0013127145,0.0028160026],"category_scores_gemma":[0.069930285,0.0011896464,0.0012386073,0.0026541941,0.0010125181,0.0046625705,0.005056864,0.0019274784,0.001651132],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025451338,0.00054528605,0.010439273,0.0008448633,0.000107686275,0.0003447522,0.005316867,0.016049594,0.044307627,0.010317446,0.0030882133,0.90838397],"study_design_scores_gemma":[0.00035787388,0.0012124974,0.04159052,0.0024309417,0.000591399,0.0026579867,0.008108013,0.43984666,0.25121668,0.07191598,0.17925307,0.0008184275],"about_ca_topic_score_codex":0.00294508,"about_ca_topic_score_gemma":0.0051281475,"teacher_disagreement_score":0.015956651,"about_ca_system_score_codex":0.0014303379,"about_ca_system_score_gemma":0.0029255995,"threshold_uncertainty_score":0.08438784},"labels":[],"label_agreement":null},{"id":"W2080703278","doi":"10.1109/icsme.2014.79","title":"Semi-automatic Identification and Representation of Subsystem Variability in Simulink Models","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Identification (biology); Set (abstract data type); Representation (politics); Process (computing); Inference; Data mining; Software; Software engineering; Artificial intelligence; Programming language","score_opus":0.02297703453982183,"score_gpt":0.2822139695306576,"score_spread":0.2592369349908358,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080703278","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03063181,0.00002220523,0.96320754,0.000034667926,0.000005484828,0.000049470622,0.00030088966,0.0053526303,0.00039541416],"genre_scores_gemma":[0.38168007,0.00006429521,0.61554295,0.000029185698,0.0000075194134,0.00019767128,0.0010250056,0.0007663901,0.00068693626],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987054,0.0003235425,0.00012251969,0.00024317663,0.0005331576,0.000072297575],"domain_scores_gemma":[0.9960194,0.0018923904,0.0007559908,0.000869484,0.000402532,0.000060266837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012195522,0.00087372767,0.00057026,0.0013180897,0.0004269138,0.0011435786,0.0012023562,0.00057993305,0.0013568037],"category_scores_gemma":[0.00607375,0.0006470602,0.0012891742,0.0004908836,0.0006918299,0.001306992,0.0012552931,0.0009794169,0.00030056134],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033602375,0.000103558676,0.0114381835,0.00035553053,0.00012698855,0.0007855165,0.0016110529,0.74307525,0.053247202,0.02963274,0.00181316,0.15747476],"study_design_scores_gemma":[0.000015787926,0.00003678793,0.00059135346,0.000024782386,0.000025921114,0.00010314607,0.000063510095,0.9674037,0.018664671,0.010704424,0.0023420034,0.000023984681],"about_ca_topic_score_codex":0.003691299,"about_ca_topic_score_gemma":0.0049536307,"teacher_disagreement_score":0.003691299,"about_ca_system_score_codex":0.00062749686,"about_ca_system_score_gemma":0.0011863105,"threshold_uncertainty_score":0.0073395967},"labels":[],"label_agreement":null},{"id":"W2081104485","doi":"10.1109/icpc.2013.6613838","title":"Insight into a method co-change pattern to identify highly coupled methods: An empirical study","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; Computer science; clone (Java method); Commit; Software; Software evolution; Software system; Source code; Data mining; Programming language; Artificial intelligence; Software construction; Biology; Database","score_opus":0.10564983106112251,"score_gpt":0.47730012486795576,"score_spread":0.37165029380683323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081104485","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9866273,0.00014829471,0.011524025,0.00012446968,0.000007867318,0.00028333315,0.00032308328,0.00013864675,0.00082297216],"genre_scores_gemma":[0.9814997,0.0000853289,0.017258123,0.000050989885,0.000008523541,0.00014675852,0.0005246716,0.000053343203,0.00037258407],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98607117,0.004567068,0.001785183,0.0026490535,0.0044227187,0.0005047761],"domain_scores_gemma":[0.66581553,0.2610566,0.030791955,0.018347997,0.021655872,0.0023320743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01334979,0.0005933614,0.00040318223,0.004577048,0.0008728166,0.0016633414,0.0013756355,0.0012073703,0.0013471222],"category_scores_gemma":[0.12692352,0.00053826923,0.0004767719,0.0032791952,0.0015950508,0.0032885766,0.001529551,0.0017428113,0.00035936746],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039475737,0.0018814534,0.8983129,0.00059644226,0.00014356365,0.0007097529,0.0099711595,0.0015050789,0.009452287,0.00080397393,0.00090426195,0.07532442],"study_design_scores_gemma":[0.00012051567,0.0014698893,0.9118839,0.0002533329,0.00021306727,0.0032476594,0.011054632,0.051355503,0.013163041,0.0014305183,0.0056895907,0.00011835003],"about_ca_topic_score_codex":0.0026256484,"about_ca_topic_score_gemma":0.0050454023,"teacher_disagreement_score":0.01334979,"about_ca_system_score_codex":0.0008285649,"about_ca_system_score_gemma":0.0010772629,"threshold_uncertainty_score":0.070601285},"labels":[],"label_agreement":null},{"id":"W2081326554","doi":"10.1016/j.infsof.2009.04.018","title":"Change impact graphs: Determining the impact of prior codechanges","year":2009,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Victoria","funders":"","keywords":"Computer science; Source code; Software; Code (set theory); Graph; Change impact analysis; Constant (computer programming); Database; Distributed computing; Software engineering; Operating system; Theoretical computer science; Programming language","score_opus":0.02094814448982087,"score_gpt":0.3012395335253905,"score_spread":0.28029138903556966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081326554","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8415891,0.0015507892,0.11258106,0.0006650089,0.00024854686,0.0006500883,0.013748641,0.009263466,0.019703366],"genre_scores_gemma":[0.93417346,0.0003873712,0.052290488,0.000104199156,0.000077960554,0.00015335255,0.009038927,0.0010449525,0.0027292098],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99684554,0.00047395282,0.00016502422,0.00045008943,0.0018592787,0.00020607599],"domain_scores_gemma":[0.9421146,0.04028098,0.004394655,0.004426866,0.0076099355,0.0011729494],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025099718,0.00070362736,0.0005889786,0.00984761,0.00066537946,0.0011575575,0.0011266522,0.0011107605,0.0047088014],"category_scores_gemma":[0.043791547,0.0004366036,0.0009999397,0.0051883664,0.00058611366,0.0033492756,0.00090425153,0.0013572547,0.0010132457],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018977759,0.0010293415,0.3941875,0.0006189374,0.00056893064,0.00066102773,0.0007215129,0.09310585,0.01196044,0.0063543776,0.015575763,0.47331858],"study_design_scores_gemma":[0.00016416097,0.0011151272,0.27920827,0.00020610711,0.00068908674,0.0010529789,0.0006597745,0.654214,0.024001447,0.022262605,0.016229238,0.00019720894],"about_ca_topic_score_codex":0.009527686,"about_ca_topic_score_gemma":0.021318842,"teacher_disagreement_score":0.00984761,"about_ca_system_score_codex":0.0007760245,"about_ca_system_score_gemma":0.0010994833,"threshold_uncertainty_score":0.018944442},"labels":[],"label_agreement":null},{"id":"W2081439463","doi":"10.1145/1082983.1083150","title":"Mining student CVS repositories for performance indicators","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Set (abstract data type); Code (set theory); Quality (philosophy); Source lines of code; Data science; Work (physics); Source code; Data mining; Software engineering; Software; Engineering; Programming language","score_opus":0.014842147429639938,"score_gpt":0.2662585567185968,"score_spread":0.2514164092889568,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081439463","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95748854,0.0007238535,0.018344197,0.00027074854,0.000054789627,0.00022387544,0.018948384,0.0019488821,0.001996804],"genre_scores_gemma":[0.936743,0.00031558017,0.024607666,0.00003833899,0.00006888676,0.00025109787,0.036085185,0.00020833428,0.0016818453],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9931375,0.0014382845,0.0009528905,0.0013444229,0.0026925628,0.00043436015],"domain_scores_gemma":[0.9504998,0.017757451,0.0090512745,0.005503552,0.014448165,0.002739818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00500648,0.00094531506,0.0010364391,0.01896901,0.0005244758,0.0025854327,0.00176851,0.0010333208,0.0009418507],"category_scores_gemma":[0.046216276,0.00042675508,0.0006192927,0.017376812,0.00033765688,0.0017857016,0.0017928684,0.0010761705,0.0011709554],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028977875,0.0007390824,0.74552983,0.00050022604,0.00020396007,0.00038780118,0.0018466042,0.0040587825,0.004185851,0.0007851735,0.008948733,0.23252417],"study_design_scores_gemma":[0.000085253116,0.0010749433,0.8575356,0.00031472428,0.00034501313,0.0010203846,0.0032892385,0.086146295,0.017678138,0.0028173588,0.029472,0.00022098514],"about_ca_topic_score_codex":0.0037523499,"about_ca_topic_score_gemma":0.0043132543,"teacher_disagreement_score":0.01896901,"about_ca_system_score_codex":0.0006558885,"about_ca_system_score_gemma":0.0016700021,"threshold_uncertainty_score":0.026477098},"labels":[],"label_agreement":null},{"id":"W2081461908","doi":"10.1109/icst.2013.24","title":"R2Fix: Automatically Generating Bug Fixes from Bug Reports","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":99,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software bug; Linux kernel; Programming language; Pointer (user interface); Software; Memory leak; Security bug; Operating system; Artificial intelligence; Memory management; Software security assurance; Cloud computing","score_opus":0.01092669478379791,"score_gpt":0.23505020531874296,"score_spread":0.22412351053494506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081461908","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08523058,0.00062558416,0.6305681,0.0003485043,0.00018053374,0.0004936744,0.003360708,0.27694383,0.0022485352],"genre_scores_gemma":[0.2295285,0.00036638958,0.74494946,0.0002070646,0.0000709316,0.00037380488,0.012530576,0.008862799,0.0031104004],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99786586,0.00045866857,0.0002022256,0.00058039173,0.0007859184,0.000106960855],"domain_scores_gemma":[0.98739874,0.006470836,0.001573878,0.0024340579,0.001889504,0.00023298961],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027774957,0.0019289267,0.0009066085,0.0044893776,0.0003451245,0.0010162203,0.0022832768,0.0013458701,0.0028277151],"category_scores_gemma":[0.017732287,0.0009620722,0.0015031592,0.0012452411,0.0006521866,0.0017566084,0.0012998773,0.0008351639,0.0017835975],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046997142,0.00046403514,0.031149408,0.0015792535,0.000386355,0.001310072,0.00095603545,0.0290026,0.034677662,0.002261047,0.051336113,0.8464075],"study_design_scores_gemma":[0.0006755243,0.0010681353,0.025587486,0.0003305516,0.00043919744,0.0035443837,0.0004184697,0.79128003,0.12099825,0.0077066966,0.04770518,0.00024611194],"about_ca_topic_score_codex":0.0021403753,"about_ca_topic_score_gemma":0.0027670958,"teacher_disagreement_score":0.0044893776,"about_ca_system_score_codex":0.00042310683,"about_ca_system_score_gemma":0.0009894756,"threshold_uncertainty_score":0.014689028},"labels":[],"label_agreement":null},{"id":"W2081509466","doi":"10.1145/1370788.1370803","title":"Multi-criteria decision analysis for customization of estimation by analogy method AQUA+","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Weighting; Analogy; Computer science; Personalization; Multiple-criteria decision analysis; Perspective (graphical); Decision tree; Quality (philosophy); Estimation; Data mining; Mathematical optimization; Artificial intelligence; Mathematics; Engineering; Systems engineering","score_opus":0.0384598081481618,"score_gpt":0.366510602226555,"score_spread":0.3280507940783932,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081509466","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011785708,0.00009339704,0.9859343,0.00012417675,0.000030934953,0.00029161296,0.000046825993,0.00023512906,0.0014579035],"genre_scores_gemma":[0.1871985,0.000078916295,0.81063515,0.00008670095,0.000028145423,0.0009866684,0.000118740325,0.00007467065,0.00079252257],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97284275,0.022301914,0.00069140206,0.00090589136,0.0029260663,0.00033197823],"domain_scores_gemma":[0.9725948,0.022660533,0.0009292727,0.0011765494,0.0023705675,0.00026829523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01968254,0.0012075929,0.002130652,0.0022929562,0.0009765984,0.0017606432,0.001564473,0.0013305569,0.0075375824],"category_scores_gemma":[0.039366543,0.00062984414,0.0018359761,0.0018518948,0.0010655853,0.0017465234,0.0025180827,0.001925149,0.0004834012],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009906992,0.00046116827,0.0036378957,0.0010897032,0.00043476716,0.0003179084,0.00054610864,0.5664043,0.0038766398,0.087407745,0.0041250424,0.33070803],"study_design_scores_gemma":[0.000053750384,0.00023198132,0.00051819097,0.0000418864,0.0000467941,0.000052023177,0.00004081524,0.9812794,0.0006347587,0.015330354,0.0017422684,0.000027749928],"about_ca_topic_score_codex":0.0024162487,"about_ca_topic_score_gemma":0.0017790884,"teacher_disagreement_score":0.01968254,"about_ca_system_score_codex":0.0016578324,"about_ca_system_score_gemma":0.0027845306,"threshold_uncertainty_score":0.10409248},"labels":[],"label_agreement":null},{"id":"W2081705453","doi":"10.1145/1321631.1321718","title":"Assisting potentially-repetitive small-scale changes via semi-automated heuristic search","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Heuristic; Codebase; Scale (ratio); Semantic change; Change detection; Code (set theory); Software; Artificial intelligence; Information retrieval; Programming language; Set (abstract data type)","score_opus":0.019958638566838938,"score_gpt":0.2753731868180923,"score_spread":0.25541454825125337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081705453","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17184709,0.0005422162,0.8109246,0.0004953989,0.00003355921,0.00091171905,0.00041870956,0.00944885,0.0053779758],"genre_scores_gemma":[0.36066887,0.00020757191,0.636343,0.00016783766,0.000020216155,0.00040924057,0.000822447,0.00029987402,0.0010609447],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972568,0.0011101578,0.00019868008,0.000548714,0.0007160902,0.00016958505],"domain_scores_gemma":[0.9793194,0.015786635,0.0014384513,0.0019020253,0.0012644556,0.00028911635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032822436,0.0015576169,0.0017945204,0.0026020145,0.001064771,0.0017009354,0.0028809933,0.001840885,0.0029992769],"category_scores_gemma":[0.01978057,0.00083760504,0.00085961336,0.0019425269,0.0010973017,0.0019867604,0.0019712248,0.0010290943,0.0008622138],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014159074,0.0013139719,0.016460622,0.0015130937,0.00034860746,0.0012082219,0.003173906,0.20340495,0.03681876,0.012419009,0.009369525,0.7125534],"study_design_scores_gemma":[0.00033351115,0.00047760783,0.002669909,0.0001415203,0.00023797787,0.0004961658,0.0012129846,0.9538395,0.014775783,0.018982524,0.0067422814,0.000090365065],"about_ca_topic_score_codex":0.005057757,"about_ca_topic_score_gemma":0.013486546,"teacher_disagreement_score":0.005057757,"about_ca_system_score_codex":0.00094471767,"about_ca_system_score_gemma":0.0029922344,"threshold_uncertainty_score":0.017358422},"labels":[],"label_agreement":null},{"id":"W2081943844","doi":"10.1016/j.cag.2005.03.007","title":"3D visualization techniques to support slicing-based program comprehension","year":2005,"lang":"en","type":"article","venue":"Computers & Graphics","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada; Concordia University","keywords":"Computer science; Program slicing; Visualization; Program comprehension; Software visualization; Slicing; Source code; Rendering (computer graphics); Software; Comprehension; Human–computer interaction; Software system; Programming language; Data mining; Component-based software engineering; Artificial intelligence; Computer graphics (images)","score_opus":0.0232504246986767,"score_gpt":0.32035949055804447,"score_spread":0.29710906585936775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081943844","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011797989,0.0001915271,0.9570017,0.00025821698,0.00008207641,0.00007005045,0.00027145116,0.027018594,0.0033083877],"genre_scores_gemma":[0.1577762,0.00052322843,0.83109134,0.00022641612,0.000083630766,0.00019783917,0.0007704651,0.005988224,0.0033426974],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995055,0.00013157609,0.00004360249,0.000056913934,0.00022233289,0.00004003723],"domain_scores_gemma":[0.9962585,0.0018880725,0.00024764193,0.00074737903,0.0007259486,0.00013241243],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00074908335,0.0015809219,0.000632509,0.0012136399,0.00044809107,0.0014276305,0.0013053966,0.00086324546,0.016803466],"category_scores_gemma":[0.0051054223,0.00077312405,0.000864089,0.00097038824,0.00042660173,0.0022376613,0.0017188878,0.0018782538,0.0014741401],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005094944,0.00029753597,0.002040175,0.0010105909,0.00014194098,0.00069373625,0.0034569674,0.031731375,0.24794242,0.038000498,0.039665688,0.63450956],"study_design_scores_gemma":[0.00024433207,0.0002865352,0.0019955414,0.0002545995,0.00018500737,0.0007317572,0.00034413324,0.6196967,0.2550193,0.036165055,0.08491421,0.00016283181],"about_ca_topic_score_codex":0.0014595337,"about_ca_topic_score_gemma":0.0029078822,"teacher_disagreement_score":0.016803466,"about_ca_system_score_codex":0.00031674327,"about_ca_system_score_gemma":0.00064228015,"threshold_uncertainty_score":0.0562132},"labels":[],"label_agreement":null},{"id":"W2081946494","doi":"10.5555/776816.776859","title":"Design pattern rationale graphs: linking design to source","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Software design pattern; Computer science; Design pattern; Structural pattern; Flexibility (engineering); Pattern language (formal languages); Source code; Software engineering; Code (set theory); TRACE (psycholinguistics); Representation (politics); Human–computer interaction; Programming language; Software design; Software development; Software","score_opus":0.05024159134126945,"score_gpt":0.2606800102843191,"score_spread":0.21043841894304963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081946494","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0023022448,0.00018497717,0.98203087,0.00059555273,0.00006609637,0.00036752314,0.0008425552,0.009206885,0.004403243],"genre_scores_gemma":[0.025677396,0.0006833629,0.962785,0.00035478492,0.00004593555,0.0007075501,0.0036203163,0.0026338624,0.003491884],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99375015,0.002692861,0.00058885344,0.0006374868,0.0021324307,0.00019806625],"domain_scores_gemma":[0.9744899,0.015078561,0.0021562728,0.005043908,0.0028515018,0.00037986608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007383467,0.0020193916,0.0006745569,0.008280721,0.0013889961,0.004174421,0.0027147247,0.002600588,0.009481423],"category_scores_gemma":[0.04209631,0.0019227244,0.0017208124,0.0055491487,0.00200124,0.0070816134,0.004639205,0.0027470551,0.0032972894],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017122453,0.00037037622,0.004426539,0.0017950195,0.00018019855,0.0013398456,0.00424583,0.028788434,0.007739172,0.23611084,0.04775683,0.66707575],"study_design_scores_gemma":[0.00015709027,0.00014049636,0.001755868,0.0011515961,0.00016052043,0.0010602309,0.0009671846,0.16793337,0.017730344,0.39108056,0.41765508,0.00020766194],"about_ca_topic_score_codex":0.0070479503,"about_ca_topic_score_gemma":0.008751123,"teacher_disagreement_score":0.009481423,"about_ca_system_score_codex":0.0014560069,"about_ca_system_score_gemma":0.0038134558,"threshold_uncertainty_score":0.039047956},"labels":[],"label_agreement":null},{"id":"W2082183626","doi":"10.1145/2393596.2393670","title":"An industrial study on the risk of software changes","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":145,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada); Polytechnique Montréal; Queen's University","funders":"","keywords":"Computer science; Software bug; Code review; Source lines of code; Reliability (semiconductor); Code (set theory); Focus (optics); Software quality; Software; Software engineering; Software development; Data science; Programming language","score_opus":0.08401407539052598,"score_gpt":0.3082452900654755,"score_spread":0.22423121467494955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2082183626","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99424225,0.00037003265,0.0035278553,0.0002081106,0.0000110749925,0.000074032,0.00007455795,0.000028636108,0.0014635326],"genre_scores_gemma":[0.99643695,0.000239139,0.0028615545,0.00008448282,0.00002177266,0.000030517593,0.00011507843,0.000008268345,0.00020232656],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.988035,0.0063607893,0.0008302638,0.0015997961,0.0028408007,0.00033330024],"domain_scores_gemma":[0.6579019,0.29069132,0.02363461,0.01021735,0.015464955,0.0020898585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015251207,0.00046357338,0.0005041442,0.003333781,0.00095032697,0.0013264933,0.0009840708,0.0010678244,0.0008764601],"category_scores_gemma":[0.11942178,0.00049519195,0.0006100345,0.00270441,0.0011531903,0.00253153,0.0012115147,0.0013534713,0.00029498778],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023512729,0.0013718595,0.93854034,0.00018168119,0.00016913954,0.00062177284,0.0069909203,0.0040980284,0.0010302899,0.00079164904,0.0008879617,0.045081228],"study_design_scores_gemma":[0.000084070736,0.0030345262,0.9440707,0.00019273258,0.00021060069,0.0011693073,0.005574613,0.038017374,0.0022136807,0.0017339144,0.003614724,0.00008370553],"about_ca_topic_score_codex":0.0038233958,"about_ca_topic_score_gemma":0.0039878553,"teacher_disagreement_score":0.015251207,"about_ca_system_score_codex":0.0012922193,"about_ca_system_score_gemma":0.0006298585,"threshold_uncertainty_score":0.080657065},"labels":[],"label_agreement":null},{"id":"W2083101347","doi":"10.1109/vissof.2007.4290707","title":"Task-specific source code dependency investigation","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Dependency (UML); Reuse; Task (project management); Identification (biology); Source code; Software engineering; Human–computer interaction; Code (set theory); Software; Code reuse; Programming language; Systems engineering; Engineering","score_opus":0.025893892716861443,"score_gpt":0.2588888448406758,"score_spread":0.23299495212381438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2083101347","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7202626,0.00024166654,0.2615207,0.00024935786,0.000057318673,0.0012543135,0.001260125,0.0064155934,0.0087382505],"genre_scores_gemma":[0.84002733,0.00013309203,0.1533369,0.00013795501,0.000016701082,0.00092142704,0.0017996818,0.00064640056,0.0029805074],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971956,0.0010316947,0.00029935455,0.00061996124,0.0007004476,0.0001530154],"domain_scores_gemma":[0.9681612,0.019469315,0.0020446063,0.00566195,0.0042637317,0.0003992086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019821697,0.0006375767,0.00046899126,0.001038217,0.00046631924,0.00057416665,0.0007212922,0.0007745661,0.0027701554],"category_scores_gemma":[0.03016183,0.00030369987,0.0003396044,0.00076108094,0.00042168802,0.0008759988,0.0013382783,0.0007736544,0.00053495105],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014455538,0.0027334725,0.039760392,0.0029002554,0.00016198182,0.001296291,0.014484187,0.010069707,0.39960542,0.005567317,0.012070196,0.5099052],"study_design_scores_gemma":[0.0005149664,0.0039654523,0.21383375,0.0003524364,0.0003911222,0.0034841993,0.004822003,0.23150292,0.45159343,0.01781612,0.07126035,0.00046322783],"about_ca_topic_score_codex":0.0018673509,"about_ca_topic_score_gemma":0.0035415562,"teacher_disagreement_score":0.0027701554,"about_ca_system_score_codex":0.00033889737,"about_ca_system_score_gemma":0.0009561313,"threshold_uncertainty_score":0.010482848},"labels":[],"label_agreement":null},{"id":"W2083587042","doi":"10.1145/1882291.1882351","title":"Using dynamic analysis to create trace-focused user interfaces for IDEs","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; TRACE (psycholinguistics); Human–computer interaction; User interface; Focus (optics); Visualization; User interface design; Software visualization; Interface (matter); Software; Graphical user interface; Task (project management); Software development; Software engineering; User experience design; Software construction; Systems engineering; Engineering; Operating system; Data mining","score_opus":0.03178156484843799,"score_gpt":0.33486833681602246,"score_spread":0.3030867719675845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2083587042","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065004798,0.00006085636,0.9855419,0.00007588879,0.000021786633,0.00005222973,0.00006966086,0.0058694803,0.001807702],"genre_scores_gemma":[0.100797206,0.00016795348,0.89322364,0.0000956534,0.000017281678,0.00018923504,0.00035077563,0.0026531585,0.0025050845],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99828917,0.0006049315,0.00011374666,0.00027912928,0.0005767318,0.00013629685],"domain_scores_gemma":[0.9936691,0.004040792,0.0002653916,0.0011382212,0.0007045956,0.00018204715],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029046435,0.0012360198,0.0005344544,0.002041315,0.0005738036,0.0029818336,0.0016012853,0.0010227939,0.00624533],"category_scores_gemma":[0.012056289,0.0007836545,0.00097286556,0.0006177212,0.0008721639,0.0042782347,0.0034769303,0.0015719075,0.0015802394],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070910476,0.00032911744,0.0034930648,0.00066960335,0.00012547313,0.0010932909,0.009165027,0.020862708,0.10865183,0.07568305,0.01281314,0.7664046],"study_design_scores_gemma":[0.00040032575,0.0005842778,0.0023616606,0.0005549578,0.0002235196,0.0019841483,0.001919515,0.43037006,0.20827895,0.0964993,0.25648546,0.0003379258],"about_ca_topic_score_codex":0.00054098497,"about_ca_topic_score_gemma":0.0008303314,"teacher_disagreement_score":0.00624533,"about_ca_system_score_codex":0.0003752704,"about_ca_system_score_gemma":0.0005975686,"threshold_uncertainty_score":0.02089274},"labels":[],"label_agreement":null},{"id":"W2083651196","doi":"10.1145/1988997.1989020","title":"An empirical analysis of a testability model for object-oriented programs","year":2011,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Testability; Computer science; Metric (unit); Unit testing; Java; Object-oriented programming; Empirical research; Reliability engineering; Software; Programming language; Engineering; Mathematics; Statistics","score_opus":0.06620077207080198,"score_gpt":0.31307960672924373,"score_spread":0.24687883465844174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2083651196","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88430834,0.00032569055,0.10965165,0.00069517497,0.000017569744,0.00026259795,0.0008177215,0.00025245757,0.0036687497],"genre_scores_gemma":[0.9920208,0.000043718665,0.0069712284,0.000033543798,0.000008229565,0.00015834841,0.00060991314,0.000029962499,0.00012415819],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98089683,0.01140611,0.00085415633,0.0017713136,0.004522742,0.0005488331],"domain_scores_gemma":[0.6393534,0.30866644,0.019759985,0.022589916,0.008235634,0.0013946911],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02603377,0.0008123894,0.0005986942,0.0037008848,0.00053659594,0.0021583284,0.001471881,0.0014412977,0.0014747152],"category_scores_gemma":[0.20364238,0.00034917437,0.0011408384,0.0029019157,0.001509453,0.003858749,0.0013009326,0.0017011209,0.0002394927],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006690554,0.0013439765,0.73447007,0.00039683565,0.00075321615,0.0003230278,0.001747233,0.15859452,0.0027035775,0.03090031,0.0019374425,0.0661607],"study_design_scores_gemma":[0.000077374534,0.0010395208,0.24455415,0.000121347744,0.00013541787,0.0005481882,0.00054701266,0.72823626,0.0012793923,0.021654708,0.0017467686,0.000059932096],"about_ca_topic_score_codex":0.0023394823,"about_ca_topic_score_gemma":0.0017644076,"teacher_disagreement_score":0.02603377,"about_ca_system_score_codex":0.0022583243,"about_ca_system_score_gemma":0.0011434524,"threshold_uncertainty_score":0.13768142},"labels":[],"label_agreement":null},{"id":"W2083998741","doi":"10.1016/j.infsof.2005.09.003","title":"A measurement framework for object-oriented software testability","year":2005,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Testability; Software engineering; Computer science; Code refactoring; Software metric; Software reliability testing; Reliability engineering; Software development; Software; Software construction; Engineering; Programming language","score_opus":0.01571495417069761,"score_gpt":0.26043312010953884,"score_spread":0.24471816593884124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2083998741","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0069070267,0.00021391806,0.9873533,0.00020192971,0.000037778558,0.00030908163,0.00025149737,0.0018171126,0.002908377],"genre_scores_gemma":[0.20300938,0.00018000421,0.7937568,0.00009907463,0.00006968524,0.0011407265,0.0007323057,0.0002563935,0.00075576204],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9744323,0.008822983,0.0032951757,0.0026235515,0.009677665,0.001148345],"domain_scores_gemma":[0.94523275,0.025181178,0.0064890794,0.008154315,0.013558167,0.0013844994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020208925,0.002370097,0.0015871262,0.011926457,0.001955909,0.0058657364,0.0037523431,0.0023237218,0.0024933622],"category_scores_gemma":[0.064919636,0.00092593837,0.0025148045,0.0057236566,0.0031774985,0.0077418466,0.004426104,0.003235628,0.0010686462],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017583367,0.00071391073,0.02104351,0.000714579,0.00040028855,0.00016348492,0.0018625545,0.039696988,0.012080207,0.5220626,0.0064956304,0.3945904],"study_design_scores_gemma":[0.00010741417,0.0013547244,0.034484677,0.00078280707,0.00044954347,0.0009964366,0.0009936314,0.45328677,0.012184268,0.46974456,0.025222566,0.0003926425],"about_ca_topic_score_codex":0.008843567,"about_ca_topic_score_gemma":0.005472261,"teacher_disagreement_score":0.020208925,"about_ca_system_score_codex":0.0032914006,"about_ca_system_score_gemma":0.004961346,"threshold_uncertainty_score":0.10687631},"labels":[],"label_agreement":null},{"id":"W2084409438","doi":"10.1016/j.jvlc.2010.03.001","title":"IssuePlayer: An extensible framework for visual assessment of issue management in software development projects","year":2010,"lang":"en","type":"article","venue":"Journal of Visual Languages & Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Dashboard; Team software process; Software development; Software engineering; Visualization; Software; Process (computing); Identification (biology); Software project management; Process management; Software development process; Knowledge management; Engineering management; Software construction; Engineering; Artificial intelligence","score_opus":0.02097988454379177,"score_gpt":0.4031775295866613,"score_spread":0.3821976450428695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084409438","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012934769,0.0001756142,0.8348222,0.00032230053,0.00009690008,0.001064855,0.002756976,0.14379126,0.0040351213],"genre_scores_gemma":[0.08562453,0.000290975,0.89372915,0.00017381807,0.000048999587,0.0011287557,0.0055954596,0.00797007,0.0054382496],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973598,0.00074477313,0.0005101031,0.00030825875,0.0008770711,0.00019994243],"domain_scores_gemma":[0.99072343,0.0051157493,0.0007821981,0.0013643709,0.0012828207,0.00073139777],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005893625,0.0018596108,0.0008458235,0.0060236556,0.00074215914,0.00496537,0.00285648,0.0017065557,0.012589761],"category_scores_gemma":[0.021689633,0.0014261337,0.0014683553,0.0017969134,0.00052857434,0.006652119,0.004701265,0.0019257831,0.002737938],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015729619,0.0011386351,0.015743198,0.0025511696,0.00038461108,0.0012544083,0.0074491887,0.023815403,0.027515136,0.036350355,0.0704148,0.8118101],"study_design_scores_gemma":[0.00083949213,0.0006345497,0.014779409,0.0017569639,0.000494924,0.0013284715,0.0020273488,0.57900333,0.05142194,0.068316884,0.27872643,0.00067035644],"about_ca_topic_score_codex":0.0063561187,"about_ca_topic_score_gemma":0.010514327,"teacher_disagreement_score":0.012589761,"about_ca_system_score_codex":0.00091527536,"about_ca_system_score_gemma":0.0019677577,"threshold_uncertainty_score":0.04211694},"labels":[],"label_agreement":null},{"id":"W2084481113","doi":"10.1109/dasc.2011.42","title":"A Natural Classification Scheme for Software Security Patterns","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software security assurance; Computer science; Security testing; Security information and event management; Computer security model; Security bug; Scheme (mathematics); Security through obscurity; Classification scheme; Security service; Secure coding; Security engineering; Software; Computer security; Cloud computing security; Information security; Machine learning; Mathematics; Operating system","score_opus":0.05391884910734808,"score_gpt":0.28115576882404103,"score_spread":0.22723691971669296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084481113","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011980532,0.00036006025,0.9681788,0.0015718156,0.00030553647,0.0011458509,0.0015494544,0.0018992071,0.013008725],"genre_scores_gemma":[0.052045245,0.00027130506,0.93937045,0.0003430478,0.0001231573,0.0009979791,0.0024307203,0.00015095895,0.0042672027],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99029166,0.0019401718,0.003280995,0.0017278313,0.0022186365,0.00054062845],"domain_scores_gemma":[0.98621494,0.0026196968,0.0018436245,0.004048826,0.004645849,0.00062703673],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005546084,0.001110781,0.00071825524,0.008262571,0.0026131333,0.004605393,0.0019737235,0.0024674418,0.005115617],"category_scores_gemma":[0.012934954,0.00066140585,0.0024056677,0.005943038,0.0035114416,0.010017628,0.0024039757,0.002978633,0.0031675387],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009376083,0.00018029989,0.0059709344,0.00041439425,0.000046851783,0.00026814055,0.0014117307,0.0025636647,0.006744766,0.8008433,0.0147306975,0.16673128],"study_design_scores_gemma":[0.00011081361,0.00040183077,0.0045627886,0.0006738062,0.0001232446,0.0032762473,0.0010959695,0.055728484,0.00701774,0.6149858,0.31180423,0.00021908685],"about_ca_topic_score_codex":0.0027908722,"about_ca_topic_score_gemma":0.0025422173,"teacher_disagreement_score":0.008262571,"about_ca_system_score_codex":0.0020934609,"about_ca_system_score_gemma":0.0032074603,"threshold_uncertainty_score":0.02933079},"labels":[],"label_agreement":null},{"id":"W2084918733","doi":"10.1142/s0218194012500118","title":"APPLYING EXPERT JUDGMENT TO IMPROVE AN INDIVIDUAL'S ABILITY TO PREDICT SOFTWARE DEVELOPMENT EFFORT","year":2012,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Consejo Nacional de Ciencia y Tecnología","keywords":"Computer science; Software; Set (abstract data type); Software engineering; Personal software process; Schedule; Process (computing); Best practice; Software bug; Software development; Data science; Engineering management; Engineering; Software construction; Management","score_opus":0.017445813949862294,"score_gpt":0.27626413262042093,"score_spread":0.25881831867055866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084918733","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9762716,0.00011053113,0.019629551,0.00017978002,0.000032420718,0.00008603642,0.00011613306,0.00019084412,0.0033829205],"genre_scores_gemma":[0.99008477,0.000044930748,0.0091697015,0.00004755872,0.000015725316,0.00002907428,0.00010619565,0.0000074714153,0.0004946041],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9972198,0.0014486758,0.00018099416,0.00054515735,0.000470411,0.00013486281],"domain_scores_gemma":[0.9643967,0.027976021,0.0019671211,0.0020916616,0.002732966,0.0008355399],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072983336,0.00048093768,0.00064784515,0.0014887673,0.00047080664,0.0010813386,0.00053232646,0.00080901245,0.0016004812],"category_scores_gemma":[0.04168973,0.0001990341,0.00039966733,0.0008487008,0.00027641156,0.0011918055,0.00057984435,0.0007743115,0.0006896184],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009833795,0.0018102047,0.721083,0.00009843503,0.00033936888,0.000087821034,0.0012331544,0.017439473,0.0030876116,0.00040128545,0.0021680598,0.25126818],"study_design_scores_gemma":[0.00018182062,0.0035386202,0.5646148,0.00009886623,0.0003126297,0.0002796485,0.0019222399,0.41329536,0.008000387,0.0056538144,0.0019705954,0.0001311711],"about_ca_topic_score_codex":0.0031017163,"about_ca_topic_score_gemma":0.004524216,"teacher_disagreement_score":0.0072983336,"about_ca_system_score_codex":0.0002881105,"about_ca_system_score_gemma":0.0005294628,"threshold_uncertainty_score":0.038597763},"labels":[],"label_agreement":null},{"id":"W2084979378","doi":"10.1145/568235.568237","title":"Dynamic analysis for reverse engineering and program understanding","year":2002,"lang":"en","type":"article","venue":"ACM SIGAPP Applied Computing Review","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Reverse engineering; Interoperation; Computer science; Software engineering; Legacy system; Program comprehension; Component (thermodynamics); Software maintenance; Systems engineering; Focus (optics); Software system; Data science; Software; Engineering; Interoperability; World Wide Web; Programming language","score_opus":0.051934204248897435,"score_gpt":0.3006018289755211,"score_spread":0.24866762472662368,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084979378","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002141305,0.0019587267,0.9849474,0.0014983587,0.00010331652,0.000109510984,0.00009807705,0.0010193707,0.008123902],"genre_scores_gemma":[0.09355113,0.0053247525,0.8886175,0.0007002961,0.00030971927,0.000578546,0.00060314254,0.001231657,0.009083233],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9925322,0.003264787,0.00056775595,0.0011289937,0.002110705,0.00039556174],"domain_scores_gemma":[0.9809749,0.01082102,0.0013946712,0.0040973686,0.0025371807,0.00017489116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006472585,0.0018813951,0.0013970102,0.0054852334,0.0016154667,0.006107364,0.0023971207,0.001816944,0.008132086],"category_scores_gemma":[0.02207774,0.0010377088,0.002375373,0.0041159526,0.005602843,0.0112784775,0.0048709083,0.004862837,0.001995529],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048823902,0.000056786208,0.0008385886,0.0006683465,0.000060385853,0.00029617912,0.0012249018,0.009726429,0.003252922,0.77395886,0.0058775134,0.2039903],"study_design_scores_gemma":[0.000021253036,0.000033813525,0.00043642533,0.00035697652,0.000058336725,0.00048239954,0.00063381315,0.06657635,0.0060436158,0.8458121,0.07948761,0.000057382804],"about_ca_topic_score_codex":0.003965286,"about_ca_topic_score_gemma":0.0030760788,"teacher_disagreement_score":0.008132086,"about_ca_system_score_codex":0.002967859,"about_ca_system_score_gemma":0.0038904164,"threshold_uncertainty_score":0.03423071},"labels":[],"label_agreement":null},{"id":"W2085597081","doi":"10.1109/tse.2013.27","title":"The Impact of Classifier Configuration and Classifier Combination on Bug Localization","year":2013,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Classifier (UML); Computer science; Artificial intelligence; Machine learning; Random subspace method; Source code; Probabilistic classification; Quadratic classifier; Data mining; Pattern recognition (psychology); Support vector machine; Naive Bayes classifier; Programming language","score_opus":0.013915283300651619,"score_gpt":0.24821500259989007,"score_spread":0.23429971929923846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2085597081","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95140713,0.010654561,0.025231441,0.0015126854,0.00068819826,0.00058282184,0.00079145556,0.0045470865,0.0045846426],"genre_scores_gemma":[0.98006964,0.0007411769,0.016451383,0.0003187121,0.00020003102,0.00019421501,0.0010778852,0.00032570615,0.00062128954],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9662212,0.014309876,0.004156333,0.00836386,0.004861466,0.0020871838],"domain_scores_gemma":[0.81041646,0.14270714,0.008972839,0.02130262,0.012273375,0.0043274667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029137626,0.0049952907,0.0037960983,0.0058547407,0.0029386324,0.005440454,0.0032457644,0.004519301,0.0013052499],"category_scores_gemma":[0.120930776,0.0019142148,0.0020257016,0.0048329155,0.0026089463,0.011590029,0.0032470978,0.0042903135,0.0014245462],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008782708,0.0029235824,0.29555288,0.0012875304,0.0031354604,0.0011441316,0.0011435627,0.21447873,0.023629395,0.00096890936,0.012393992,0.43455905],"study_design_scores_gemma":[0.0014538725,0.015429636,0.14513414,0.0008166039,0.006380394,0.0055744858,0.0036038328,0.72750235,0.07174043,0.008071773,0.013088311,0.0012041584],"about_ca_topic_score_codex":0.0043016137,"about_ca_topic_score_gemma":0.0037196036,"teacher_disagreement_score":0.029137626,"about_ca_system_score_codex":0.0019843348,"about_ca_system_score_gemma":0.0022641628,"threshold_uncertainty_score":0.15409636},"labels":[],"label_agreement":null},{"id":"W2085689804","doi":"10.1142/s0218194009004167","title":"A FRAMEWORK FOR TOOL-BASED SOFTWARE ARCHITECTURE RECONSTRUCTION","year":2009,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Reference architecture; Software architecture description; Architecture tradeoff analysis method; Resource-oriented architecture; Computer science; Software architecture; Multilayered architecture; Software engineering; Applications architecture; Space-based architecture; Architecture; Database-centric architecture; Software; Computer architecture; Software construction; Software system; Programming language","score_opus":0.01110723504304779,"score_gpt":0.26480299377069977,"score_spread":0.253695758727652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2085689804","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00036717168,0.00006144285,0.9969072,0.0001034415,0.000009318369,0.00012676674,0.00004649854,0.0014617762,0.0009164632],"genre_scores_gemma":[0.010120818,0.00012572981,0.9883454,0.000036476646,0.000010409514,0.00017976912,0.0003445307,0.00017531078,0.0006616224],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99154496,0.0024170487,0.0010803291,0.0013964466,0.003032659,0.0005285247],"domain_scores_gemma":[0.99199873,0.0025478404,0.00059530756,0.0031451534,0.0014476324,0.00026525266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011136554,0.0020397953,0.0013786402,0.006181892,0.0023009048,0.0071455315,0.0063660704,0.0035533821,0.005091587],"category_scores_gemma":[0.014714158,0.0021215251,0.0054894667,0.003937716,0.005596858,0.008326087,0.006607852,0.0054818187,0.002480897],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000052083447,0.00015426386,0.0009629032,0.00052729365,0.00013779415,0.000850707,0.0015647084,0.057504654,0.004360542,0.77126765,0.004759683,0.15785755],"study_design_scores_gemma":[0.00005968607,0.00013541422,0.00051248766,0.00065112073,0.00012816128,0.0010679424,0.0005690689,0.36383876,0.007926013,0.50264174,0.12230909,0.00016053223],"about_ca_topic_score_codex":0.008299405,"about_ca_topic_score_gemma":0.0073690545,"teacher_disagreement_score":0.011136554,"about_ca_system_score_codex":0.0028802892,"about_ca_system_score_gemma":0.004945359,"threshold_uncertainty_score":0.058896363},"labels":[],"label_agreement":null},{"id":"W2085695897","doi":"10.1109/icsm.2013.102","title":"Interactive Exploration of Collaborative Software-Development Data","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Visualization; Data visualization; Data exploration; Software; Data science; Work (physics); Software development; Software engineering; Data mining; Engineering","score_opus":0.055043804195719234,"score_gpt":0.3072417476859928,"score_spread":0.2521979434902736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2085695897","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29670465,0.0008627715,0.65978754,0.0012737891,0.00014457568,0.0003868544,0.0066404077,0.014773166,0.019426141],"genre_scores_gemma":[0.6092617,0.00048419926,0.3835541,0.00011112416,0.00006556938,0.00032741635,0.0030307397,0.0011495188,0.0020155904],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99812037,0.00081218325,0.000115703384,0.00022182183,0.0006002587,0.00012965074],"domain_scores_gemma":[0.98534805,0.011847223,0.0004215619,0.0012016044,0.00072427973,0.0004573088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036102554,0.00089112035,0.0008659137,0.0043241545,0.001160161,0.004126646,0.0010470394,0.0010319857,0.0052193264],"category_scores_gemma":[0.0109418575,0.000406551,0.0007375853,0.003258913,0.0008913951,0.0026168472,0.0053606615,0.0012747209,0.0005702344],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029609567,0.0008196311,0.03321051,0.004201214,0.00065115426,0.004012049,0.10563399,0.0736389,0.105037615,0.0650379,0.03521079,0.5695853],"study_design_scores_gemma":[0.00044571067,0.00054811296,0.054417174,0.001202142,0.0003269198,0.0027023368,0.021528,0.52274185,0.07095926,0.13168885,0.19278546,0.00065418665],"about_ca_topic_score_codex":0.0019864957,"about_ca_topic_score_gemma":0.0031337808,"teacher_disagreement_score":0.0052193264,"about_ca_system_score_codex":0.0005867394,"about_ca_system_score_gemma":0.0008341038,"threshold_uncertainty_score":0.019093096},"labels":[],"label_agreement":null},{"id":"W2085996689","doi":"10.1145/2024445.2024458","title":"Causes of premature aging during software development","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Artifact (error); Computer science; GRASP; Categorization; Software maintenance; Software engineering; Software development; Work (physics); Software; Code (set theory); Architecture; Risk analysis (engineering); Artificial intelligence; Programming language; Engineering; Business","score_opus":0.028211288975703616,"score_gpt":0.24003106135109425,"score_spread":0.21181977237539062,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2085996689","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96079093,0.009159439,0.012311605,0.0028629298,0.00021499963,0.00015373493,0.00022867632,0.00036178553,0.013915869],"genre_scores_gemma":[0.99324715,0.0017451888,0.002787255,0.00030680705,0.00007893047,0.000039840535,0.00012589966,0.00006552083,0.0016034085],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9933565,0.0018781335,0.0007987565,0.0008327259,0.0025154767,0.00061847124],"domain_scores_gemma":[0.9364211,0.021631423,0.021067727,0.0055049714,0.013229161,0.0021455719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006236121,0.00043185003,0.00043309736,0.0026529334,0.0016698042,0.0013666864,0.001001878,0.001184188,0.0023735983],"category_scores_gemma":[0.04649987,0.00043464522,0.00052822055,0.0020141322,0.0012236341,0.0020344858,0.0022019383,0.0011342549,0.00056285097],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024351185,0.00016152338,0.7268997,0.0008779123,0.000118340824,0.007931492,0.024090726,0.0020863514,0.004358071,0.0056619365,0.0038796195,0.22369082],"study_design_scores_gemma":[0.000043896795,0.0007031148,0.8764299,0.0008966607,0.0003195335,0.020726783,0.016410725,0.00544611,0.0070847557,0.01808381,0.05372097,0.00013364507],"about_ca_topic_score_codex":0.0035345005,"about_ca_topic_score_gemma":0.0033485596,"teacher_disagreement_score":0.006236121,"about_ca_system_score_codex":0.001627418,"about_ca_system_score_gemma":0.0018821992,"threshold_uncertainty_score":0.032980144},"labels":[],"label_agreement":null},{"id":"W2086095861","doi":"10.1109/tr.2012.2183912","title":"Evaluating Stratification Alternatives to Improve Software Defect Prediction","year":2012,"lang":"en","type":"article","venue":"IEEE Transactions on Reliability","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Machine learning; Software bug; Skewness; Software; Artificial intelligence; Software quality; Data mining; Sampling (signal processing); Stratification (seeds); Reliability engineering; Statistics; Software development; Engineering; Mathematics","score_opus":0.04061733545592563,"score_gpt":0.3384401637331717,"score_spread":0.2978228282772461,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2086095861","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6962646,0.0013102575,0.29845187,0.0005228806,0.00013066868,0.0005391834,0.00022518174,0.00091556244,0.0016398507],"genre_scores_gemma":[0.93139756,0.0001465794,0.06758367,0.000108546636,0.00004807036,0.00018714173,0.00031978017,0.000029478599,0.00017906167],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9822818,0.013689081,0.00071338034,0.0012236818,0.0016102643,0.00048175847],"domain_scores_gemma":[0.8757462,0.10186068,0.006656523,0.0073182345,0.006984371,0.0014339925],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03492776,0.0013842606,0.002125474,0.0021565598,0.0007228075,0.0015906253,0.00086366874,0.0012051493,0.0011070358],"category_scores_gemma":[0.09440004,0.0004917706,0.0012884291,0.0011781675,0.0014046397,0.0027669142,0.0018251428,0.001347379,0.00026109224],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007174446,0.0021956754,0.17444998,0.0004149582,0.0010398525,0.000180394,0.001148616,0.30285016,0.013643638,0.018371204,0.0026776532,0.47585344],"study_design_scores_gemma":[0.00035866254,0.0034144681,0.030824078,0.00008936522,0.0003485269,0.000055282027,0.00026627924,0.9347515,0.007213984,0.021379497,0.0011967361,0.00010166175],"about_ca_topic_score_codex":0.0015448357,"about_ca_topic_score_gemma":0.0015006441,"teacher_disagreement_score":0.03492776,"about_ca_system_score_codex":0.0013380738,"about_ca_system_score_gemma":0.0015050073,"threshold_uncertainty_score":0.18471783},"labels":[],"label_agreement":null},{"id":"W2086221358","doi":"10.1109/52.903173","title":"Improving subjective estimates using paired comparisons","year":2001,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ericsson (Canada); Research Canada","funders":"","keywords":"Computer science; Software; Scale (ratio); Sizing; Data science; Software technical review; Software development; Software engineering; Industrial engineering; Software quality; Engineering","score_opus":0.04136745351365596,"score_gpt":0.28974969134457945,"score_spread":0.2483822378309235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2086221358","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022600912,0.0005114374,0.96851057,0.00027436274,0.0006539413,0.00080217666,0.00036565686,0.00088748,0.005393477],"genre_scores_gemma":[0.32554817,0.0004584416,0.66509545,0.0005038752,0.0007046395,0.003981701,0.0010362578,0.0005737563,0.002097691],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.86410415,0.09982403,0.006630758,0.012360531,0.016120683,0.0009598743],"domain_scores_gemma":[0.48121166,0.42628637,0.016908145,0.035389263,0.039007578,0.0011970178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.09340024,0.002365317,0.0025191908,0.0044596028,0.0015578921,0.0042084577,0.002517969,0.0019391784,0.009920135],"category_scores_gemma":[0.40868485,0.0011259124,0.0017587553,0.0033748175,0.00214774,0.0062039257,0.004913685,0.003061875,0.0021452438],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039601782,0.000920341,0.022316905,0.0033222598,0.0024045308,0.0004021163,0.0063969046,0.015663361,0.007411744,0.03690362,0.012470795,0.8878272],"study_design_scores_gemma":[0.0023918322,0.016032048,0.10878852,0.002319569,0.0041660005,0.0021545258,0.008913321,0.29858804,0.06460621,0.4213043,0.06917494,0.0015607283],"about_ca_topic_score_codex":0.00052673154,"about_ca_topic_score_gemma":0.00068702205,"teacher_disagreement_score":0.09340024,"about_ca_system_score_codex":0.00085639587,"about_ca_system_score_gemma":0.0010816117,"threshold_uncertainty_score":0.4939536},"labels":[],"label_agreement":null},{"id":"W2086335328","doi":"10.1109/scam.2011.21","title":"A Constraint Programming Approach to Conflict-Aware Optimal Scheduling of Prioritized Code Clone Refactoring","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Constraint programming; Programming language; Scheduling (production processes); Mathematical optimization; Mathematics; Software; Stochastic programming","score_opus":0.08039188590806653,"score_gpt":0.29414768304701333,"score_spread":0.2137557971389468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2086335328","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0044564377,0.0002243092,0.99187535,0.00031099896,0.00004577337,0.00013106344,0.00012212798,0.000114732335,0.002719222],"genre_scores_gemma":[0.12086523,0.00061518187,0.874367,0.00020137295,0.000080839876,0.0006229213,0.00029074436,0.00017051082,0.002786222],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979067,0.00087713456,0.000098915916,0.0003237203,0.00053908926,0.00025434652],"domain_scores_gemma":[0.99633527,0.002599605,0.00031548622,0.000120664554,0.00046701808,0.00016198787],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032217628,0.0016897589,0.0016846257,0.00137494,0.00093389413,0.0022234365,0.0031876143,0.0014055589,0.0036107309],"category_scores_gemma":[0.0069026155,0.0015607278,0.0017274101,0.0027698446,0.0012324231,0.001687002,0.0012776386,0.0026048026,0.0003490963],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034102326,0.000052127692,0.00017482025,0.00009277818,0.000038677175,0.0000787933,0.000063943146,0.96099234,0.0006480245,0.024782127,0.00093331456,0.012108921],"study_design_scores_gemma":[0.000016852518,0.000017430908,0.000048108777,0.000011470521,0.000011660025,0.000011776483,0.000014741649,0.99200636,0.00023929692,0.0067935903,0.00082033756,0.000008404121],"about_ca_topic_score_codex":0.02679409,"about_ca_topic_score_gemma":0.022763558,"teacher_disagreement_score":0.02679409,"about_ca_system_score_codex":0.0028474687,"about_ca_system_score_gemma":0.004695871,"threshold_uncertainty_score":0.05327624},"labels":[],"label_agreement":null},{"id":"W2086580201","doi":"10.1109/vissof.2005.1684320","title":"Applying Code Analysis and 3D Design Pattern Grouping to Facilitate Program Comprehension","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Program comprehension; Source code; Programming language; Unified Modeling Language; Reverse engineering; Static program analysis; Software engineering; Animation; Software design pattern; Software; Software development; Software system; Computer graphics (images)","score_opus":0.0748341516401633,"score_gpt":0.29960521124812833,"score_spread":0.22477105960796503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2086580201","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00964087,0.000037702765,0.98096013,0.00021079795,0.000016907776,0.0001043043,0.00008202452,0.0071997913,0.0017475232],"genre_scores_gemma":[0.051333793,0.00008644926,0.9461604,0.00008179194,0.0000129968985,0.00016276466,0.00023310471,0.00081798894,0.0011106738],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982242,0.00061252463,0.00017084123,0.0003261736,0.00059185835,0.000074332194],"domain_scores_gemma":[0.98923993,0.005858101,0.0010482994,0.0021700663,0.0015196534,0.000164044],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018650406,0.0013574756,0.0007727631,0.0047859144,0.0007307119,0.0023901535,0.0012276728,0.0013324439,0.0065259673],"category_scores_gemma":[0.0147980405,0.00072319916,0.0011294057,0.0021715302,0.0012189321,0.0028922998,0.0022595173,0.0011107383,0.0021669413],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017874422,0.00026761487,0.0042705536,0.00064633816,0.00007524591,0.0006639019,0.006762715,0.018337717,0.1175097,0.027165309,0.008148117,0.81597394],"study_design_scores_gemma":[0.00013960768,0.0002794987,0.005142702,0.00032677455,0.000104265775,0.0014046284,0.0014716478,0.6437864,0.18655184,0.08294457,0.07765132,0.00019678172],"about_ca_topic_score_codex":0.001504791,"about_ca_topic_score_gemma":0.002096963,"teacher_disagreement_score":0.0065259673,"about_ca_system_score_codex":0.0006341276,"about_ca_system_score_gemma":0.0011559944,"threshold_uncertainty_score":0.021831512},"labels":[],"label_agreement":null},{"id":"W2087103218","doi":"10.1002/stvr.295","title":"Eight maxims for software inspectors","year":2004,"lang":"en","type":"article","venue":"Software Testing Verification and Reliability","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Cover (algebra); Set (abstract data type); Software; Engineering; Engineering ethics; Software inspection; Computer science; Software engineering; Management science; Software development; Software quality; Mechanical engineering","score_opus":0.027674755534598493,"score_gpt":0.2663009924304416,"score_spread":0.23862623689584309,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087103218","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41460073,0.0053597568,0.08686528,0.36885592,0.008482774,0.0012471373,0.00026186477,0.0010283282,0.11329813],"genre_scores_gemma":[0.88761246,0.0015424051,0.061435606,0.015878728,0.001095999,0.00076068117,0.0002054803,0.00030582683,0.031162774],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.95807534,0.025160061,0.0018576377,0.0017728424,0.010401082,0.0027329994],"domain_scores_gemma":[0.93964684,0.029268987,0.005105658,0.0034591183,0.0147668645,0.007752481],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035414256,0.0009591367,0.0006887206,0.0017361244,0.006588692,0.00840077,0.0019541848,0.005078459,0.0035154682],"category_scores_gemma":[0.06648682,0.0007391477,0.0006764373,0.0014140205,0.009756284,0.0074363225,0.010510007,0.012835501,0.0011710202],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045583775,0.00043578213,0.011273676,0.0014300621,0.000066848545,0.0011329058,0.37141737,0.0009873825,0.014794539,0.23034765,0.14261681,0.225041],"study_design_scores_gemma":[0.00005352048,0.00057184254,0.011391728,0.0016821345,0.00003133688,0.0010339426,0.40896744,0.003323673,0.0024187374,0.14287305,0.42743123,0.00022135398],"about_ca_topic_score_codex":0.00053900905,"about_ca_topic_score_gemma":0.0007658931,"teacher_disagreement_score":0.035414256,"about_ca_system_score_codex":0.006582958,"about_ca_system_score_gemma":0.005406913,"threshold_uncertainty_score":0.18729079},"labels":[],"label_agreement":null},{"id":"W2087145202","doi":"10.1049/ip-sen:20050012","title":"Metarule-guided association rule mining for program understanding","year":2005,"lang":"en","type":"article","venue":"IEE Proceedings - Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Russian Science Foundation; University of Victoria","keywords":"Association rule learning; Computer science; Source code; Documentation; Association (psychology); Software system; Legacy system; Software; Software evolution; Data mining; Software engineering; Process (computing); Software maintenance; Unavailability; Software construction; Programming language; Reliability engineering; Engineering","score_opus":0.061344443260858396,"score_gpt":0.31215586744306517,"score_spread":0.25081142418220675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087145202","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017168108,0.00082314387,0.976718,0.00044149582,0.00003906776,0.0002824276,0.0007112349,0.0028193311,0.0009971979],"genre_scores_gemma":[0.096845366,0.00037100722,0.9001714,0.00016499056,0.00004139632,0.00040717257,0.0014690997,0.00010839369,0.0004212579],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9924596,0.0034424728,0.00090349856,0.0010875571,0.0018999014,0.00020693253],"domain_scores_gemma":[0.97341263,0.0204377,0.0019219083,0.0020582974,0.0019331821,0.00023634857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008767162,0.0010923159,0.0024214436,0.007338806,0.000943071,0.0025963613,0.002987086,0.0018744324,0.0020529374],"category_scores_gemma":[0.034705397,0.0007340592,0.0021080922,0.0056299474,0.0010454595,0.0022577937,0.0013857023,0.0020964802,0.0013299804],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006458604,0.0007362667,0.01572909,0.0011574698,0.0007852463,0.0010133014,0.0009891507,0.10830372,0.010407279,0.022135662,0.008216783,0.82988006],"study_design_scores_gemma":[0.00010553338,0.00017913422,0.003327108,0.00022346828,0.00017832711,0.00087025034,0.00023759056,0.91895425,0.0066503137,0.061063256,0.00812371,0.000086967244],"about_ca_topic_score_codex":0.004058314,"about_ca_topic_score_gemma":0.0048231855,"teacher_disagreement_score":0.008767162,"about_ca_system_score_codex":0.0009508899,"about_ca_system_score_gemma":0.0028010642,"threshold_uncertainty_score":0.046365738},"labels":[],"label_agreement":null},{"id":"W2087156818","doi":"10.1145/1082983.1083302","title":"Quality, cleanroom and formal methods","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cleanroom; Computer science; Software engineering; Software development; Formal methods; Quality (philosophy); Systems engineering; Software; Reliability engineering; Engineering; Programming language","score_opus":0.035882057953362016,"score_gpt":0.34350678952848585,"score_spread":0.3076247315751238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087156818","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019312755,0.0046668528,0.98056763,0.0043334393,0.0002604036,0.00008120403,0.000032253418,0.0005750109,0.0075519527],"genre_scores_gemma":[0.088613234,0.006111497,0.89274997,0.0015130739,0.0006452679,0.00033247966,0.00014193024,0.00041782946,0.009474757],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98590547,0.005130225,0.0011624439,0.0015846859,0.005560196,0.0006570192],"domain_scores_gemma":[0.9685148,0.019513974,0.002227985,0.0052330275,0.0039629573,0.00054735405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015922403,0.00145184,0.001200221,0.0043115057,0.0015766586,0.0052760807,0.0037770045,0.0022386059,0.0043351743],"category_scores_gemma":[0.025672413,0.0010989233,0.002871777,0.0020523681,0.016767753,0.013639793,0.004812652,0.0059913946,0.0010337643],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020807052,0.00004757643,0.00029764863,0.00046167683,0.000033899676,0.00008607989,0.00054006255,0.008810828,0.0007284794,0.919305,0.002172629,0.06749535],"study_design_scores_gemma":[0.00004727976,0.00007034131,0.00010891122,0.00029708256,0.00003785538,0.00017565474,0.00013258622,0.014825189,0.0015027963,0.9253752,0.057370994,0.000056024368],"about_ca_topic_score_codex":0.0037055537,"about_ca_topic_score_gemma":0.0023412928,"teacher_disagreement_score":0.015922403,"about_ca_system_score_codex":0.0040051467,"about_ca_system_score_gemma":0.004241285,"threshold_uncertainty_score":0.0842067},"labels":[],"label_agreement":null},{"id":"W2087339015","doi":"10.1145/1111449.1111491","title":"Who's asking for help?","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Personalization; User modeling; Human–computer interaction; Bayesian network; Scope (computer science); Probabilistic logic; Usability; Task (project management); Intrusiveness; Software; User interface; Dynamic Bayesian network; Machine learning; Artificial intelligence; World Wide Web","score_opus":0.015236640409315133,"score_gpt":0.2622404815263549,"score_spread":0.24700384111703974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087339015","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8024947,0.0028650975,0.12332443,0.030564617,0.00032239725,0.00034772579,0.0011381669,0.0051721265,0.033770733],"genre_scores_gemma":[0.9677921,0.0006733098,0.02750302,0.000943806,0.00007330037,0.000044771477,0.00020104184,0.000093628405,0.0026751247],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99780136,0.0013229431,0.000107862106,0.00031390696,0.00026912053,0.0001846677],"domain_scores_gemma":[0.9876967,0.009070115,0.0009886832,0.0006545581,0.00080040394,0.0007895493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029420948,0.0003857608,0.00036669776,0.00063296704,0.0010789991,0.0015734304,0.0007281528,0.0019133658,0.005708808],"category_scores_gemma":[0.022408212,0.0003093881,0.00032035355,0.0007661113,0.00079916086,0.0019297705,0.0005581314,0.00082798523,0.0023950909],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002612886,0.0010073651,0.30329952,0.0008930544,0.0001587277,0.0012764793,0.009248876,0.010578395,0.016587036,0.010663554,0.037369538,0.60630465],"study_design_scores_gemma":[0.00057707727,0.0018172212,0.22649442,0.00075102015,0.000603226,0.011009901,0.03125962,0.43908682,0.028416878,0.09128016,0.16804744,0.00065624755],"about_ca_topic_score_codex":0.004959491,"about_ca_topic_score_gemma":0.009117083,"teacher_disagreement_score":0.005708808,"about_ca_system_score_codex":0.00062942883,"about_ca_system_score_gemma":0.0009980629,"threshold_uncertainty_score":0.019097865},"labels":[],"label_agreement":null},{"id":"W2087450238","doi":"10.1109/wcre.2012.56","title":"SMURF: A SVM-based Incremental Anti-pattern Detection Approach","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":114,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"","keywords":"Computer science; Support vector machine; Machine learning; Precision and recall; Artificial intelligence; Recall; Data mining; Code (set theory)","score_opus":0.02353135440962072,"score_gpt":0.25201004947464517,"score_spread":0.22847869506502444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087450238","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06722151,0.0011566634,0.9127372,0.0005220344,0.00021135426,0.00052338326,0.0010589081,0.014811956,0.0017570185],"genre_scores_gemma":[0.3697017,0.0002913843,0.62409127,0.00034986963,0.00017052732,0.00047122824,0.0024277384,0.0004022442,0.002094075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99612576,0.0008373887,0.00038460782,0.0007901916,0.0015156566,0.00034632054],"domain_scores_gemma":[0.9863949,0.007177153,0.0012314228,0.00096722424,0.0039653704,0.00026396135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045362026,0.0020963936,0.0024806869,0.0064912857,0.00083781296,0.0011952906,0.004037217,0.0022796083,0.0022239962],"category_scores_gemma":[0.017708557,0.0006508249,0.0012561232,0.002611173,0.0005448545,0.0024340497,0.0011007274,0.0021704568,0.0014662594],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038118367,0.0003290085,0.005524731,0.00026862687,0.00014689109,0.00022237247,0.00016445798,0.030357607,0.007713996,0.00091119687,0.010832553,0.94314736],"study_design_scores_gemma":[0.0000364767,0.00018423339,0.0023038446,0.000025017,0.000037041973,0.0002140745,0.000050745035,0.98672694,0.0055935234,0.0022413281,0.0025517982,0.000034964803],"about_ca_topic_score_codex":0.008726579,"about_ca_topic_score_gemma":0.008757906,"teacher_disagreement_score":0.008726579,"about_ca_system_score_codex":0.0010354212,"about_ca_system_score_gemma":0.0012992227,"threshold_uncertainty_score":0.023989975},"labels":[],"label_agreement":null},{"id":"W2087472619","doi":"10.1109/iwsm-mensura.2011.21","title":"Bidirectional Influence of Defects and Functional Size","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Functional requirement; Computer science; Measure (data warehouse); Value (mathematics); Data mining; Machine learning; Software engineering","score_opus":0.02384145946745158,"score_gpt":0.2271931709935922,"score_spread":0.20335171152614062,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087472619","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98268694,0.0010508215,0.008644522,0.00019949979,0.000019690055,0.000038079743,0.00024594116,0.00020305389,0.006911316],"genre_scores_gemma":[0.99782526,0.00019993162,0.0011284616,0.000027732181,0.000009138773,0.000009107069,0.00010200451,0.00008710857,0.00061124656],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99232274,0.0036026945,0.00024337365,0.0010742184,0.0023908839,0.0003661794],"domain_scores_gemma":[0.74153966,0.23067287,0.011625428,0.0071585346,0.0066676564,0.0023359014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037887767,0.0007928533,0.00066612783,0.0022589099,0.00028339282,0.0014503885,0.0006866598,0.00068719726,0.0043782815],"category_scores_gemma":[0.06163575,0.0005273148,0.00065384206,0.0008668581,0.0011330983,0.0013310448,0.001459469,0.0012743997,0.00064457423],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0053858953,0.00086531596,0.71136856,0.0005576052,0.0007535028,0.0014468172,0.0017442994,0.018844126,0.111354046,0.0025127016,0.0006432442,0.14452398],"study_design_scores_gemma":[0.00009139716,0.0023083792,0.92919564,0.00011258276,0.00095573865,0.0017093103,0.0008883339,0.028076166,0.030622378,0.003284517,0.002626431,0.00012906236],"about_ca_topic_score_codex":0.0022887732,"about_ca_topic_score_gemma":0.002192054,"teacher_disagreement_score":0.0043782815,"about_ca_system_score_codex":0.000659094,"about_ca_system_score_gemma":0.00059470977,"threshold_uncertainty_score":0.020037234},"labels":[],"label_agreement":null},{"id":"W2087515886","doi":"10.5555/2664446.2664477","title":"A qualitative study on performance bugs","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Software bug; Computer science; Context (archaeology); Software; Sample (material); Code (set theory); Software engineering; Operating system; Programming language","score_opus":0.07248724276815446,"score_gpt":0.39292835003969184,"score_spread":0.3204411072715374,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087515886","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9926717,0.00026385876,0.00213208,0.0016938612,0.000038666632,0.00023369722,0.00021911884,0.000024434112,0.0027225218],"genre_scores_gemma":[0.9968855,0.00021790489,0.0007279131,0.000496195,0.000011492187,0.00030713723,0.000058290938,0.000023248642,0.0012722997],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.97328067,0.020336747,0.00080262416,0.001023683,0.0027180589,0.00183811],"domain_scores_gemma":[0.8612503,0.108077794,0.010423334,0.0024612949,0.013756839,0.0040304246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023372045,0.00051552756,0.0007601177,0.0030315712,0.00496016,0.002981091,0.001492195,0.0015722143,0.0038101496],"category_scores_gemma":[0.084464066,0.00063441566,0.000353219,0.0019036551,0.007828338,0.0035520263,0.003973647,0.0021834862,0.0004616712],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000070609094,0.00011028105,0.015318552,0.0003646576,0.0000068171526,0.00048453035,0.97496337,0.00005512431,0.0011721319,0.0011371954,0.0007162861,0.0056005213],"study_design_scores_gemma":[0.000009392142,0.00019067015,0.0071601374,0.00023314654,0.0000045198617,0.00018758845,0.98478633,0.00019430566,0.0005835995,0.0003670139,0.0062587657,0.000024460489],"about_ca_topic_score_codex":0.0058480087,"about_ca_topic_score_gemma":0.0068063997,"teacher_disagreement_score":0.023372045,"about_ca_system_score_codex":0.0060002585,"about_ca_system_score_gemma":0.0035403883,"threshold_uncertainty_score":0.123604655},"labels":[],"label_agreement":null},{"id":"W2087670325","doi":"10.1007/s10664-014-9333-9","title":"Improving bug management using correlations in crash reports","year":2014,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Crash; Computer science; Eclipse; Precision and recall; Identification (biology); Software bug; Software; Data mining; Artificial intelligence; Programming language","score_opus":0.019100434507084794,"score_gpt":0.2706619834027131,"score_spread":0.25156154889562826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087670325","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8861395,0.0017031514,0.10059976,0.00077263766,0.00020200286,0.00022309176,0.0020734675,0.0050206254,0.0032658041],"genre_scores_gemma":[0.9739226,0.00022983101,0.024063561,0.000048189577,0.00008437805,0.000054033317,0.0011055691,0.00012610995,0.00036574635],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99224687,0.0029354289,0.0007781533,0.001421268,0.0022321432,0.0003860571],"domain_scores_gemma":[0.863228,0.0720521,0.034564152,0.010034559,0.017743243,0.0023779462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063381977,0.0011649611,0.000981299,0.0095089935,0.0006593525,0.0019765613,0.0011451757,0.00095307844,0.0014938563],"category_scores_gemma":[0.105690345,0.0007408197,0.00076024345,0.0053568594,0.0004091624,0.0032567268,0.0011633864,0.0014288784,0.00078984944],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059626636,0.00075260305,0.7294303,0.00033673068,0.000478531,0.00020985573,0.00051808293,0.023633568,0.0056733014,0.001290148,0.006295149,0.23078535],"study_design_scores_gemma":[0.00015769099,0.0017469391,0.48210624,0.0002748009,0.000842472,0.0006027311,0.0007632187,0.49045667,0.012550717,0.0061450345,0.0041795317,0.00017400736],"about_ca_topic_score_codex":0.006217278,"about_ca_topic_score_gemma":0.010267488,"teacher_disagreement_score":0.0095089935,"about_ca_system_score_codex":0.0008428126,"about_ca_system_score_gemma":0.0022608333,"threshold_uncertainty_score":0.033519983},"labels":[],"label_agreement":null},{"id":"W2087832061","doi":"10.1109/msr.2013.6624006","title":"Detecting API usage obstacles: A study of iOS and Android developer questions","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Android (operating system); Computer science; World Wide Web; Java; Application programming interface; Software; Software engineering; Operating system","score_opus":0.020024965077647918,"score_gpt":0.26173513519921104,"score_spread":0.24171017012156312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2087832061","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99870956,0.00005224752,0.0003665082,0.0001570967,0.0000031608154,0.000042857017,0.000055986464,0.000006743169,0.00060581136],"genre_scores_gemma":[0.9982084,0.00010454101,0.00074098696,0.00015471644,0.000010477161,0.00013592614,0.00015748155,0.000021563445,0.00046604435],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99452895,0.0026219855,0.00049928774,0.00067942997,0.0010123521,0.00065803743],"domain_scores_gemma":[0.8999589,0.07877023,0.012117282,0.0016500488,0.0051959185,0.002307583],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0080527,0.00036933814,0.0005745145,0.002783465,0.0019908273,0.0021256104,0.0009050813,0.001521402,0.0011194145],"category_scores_gemma":[0.059138034,0.00061168213,0.00037283474,0.0018868131,0.0018416449,0.0044605955,0.0024937203,0.0016347318,0.00034320756],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012886553,0.0003375641,0.26539052,0.0002420777,0.000020026473,0.0006993443,0.7142755,0.00007734531,0.0023546924,0.00049668015,0.0006845253,0.015292747],"study_design_scores_gemma":[0.000019245052,0.0003285258,0.3857062,0.00014055945,0.000023039289,0.00062007626,0.6037564,0.0015812459,0.0006752177,0.00042475492,0.0066680824,0.000056686673],"about_ca_topic_score_codex":0.008319656,"about_ca_topic_score_gemma":0.01404302,"teacher_disagreement_score":0.008319656,"about_ca_system_score_codex":0.0016427705,"about_ca_system_score_gemma":0.0012816363,"threshold_uncertainty_score":0.04258728},"labels":[],"label_agreement":null},{"id":"W2088067768","doi":"10.1109/icsme.2014.116","title":"Studying the Impact of Developer Communication on the Quality and Evolution of a Software System: A Doctoral Dissertation Retrospective","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software development; Computer science; Software peer review; Software quality; Software engineering; Software construction; Personal software process; Social software engineering; Software development process; Team software process; Package development process; Software evolution; Software analytics; Software system; Software; Programming language","score_opus":0.04751841137502348,"score_gpt":0.3508662448966308,"score_spread":0.3033478335216073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088067768","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.904699,0.01646058,0.010352456,0.00992505,0.0003867755,0.00017175799,0.00017902073,0.000046076497,0.057779316],"genre_scores_gemma":[0.95596474,0.018383184,0.004818713,0.0007802456,0.00028533163,0.00017470923,0.00019921787,0.00006479295,0.0193291],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9966157,0.0015584754,0.00014263284,0.0004612507,0.0010143098,0.00020752262],"domain_scores_gemma":[0.96743894,0.019955244,0.0028714922,0.0011921992,0.0067810337,0.0017610205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071844053,0.00026338344,0.00032444295,0.0014163596,0.0012391973,0.0035541856,0.00041616627,0.0007039097,0.0037659719],"category_scores_gemma":[0.02528724,0.00029777075,0.00037861534,0.0016043357,0.0010275436,0.0019443072,0.0014623599,0.0015306664,0.00095751975],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064271153,0.0018709475,0.13391934,0.0023377615,0.0002212165,0.00066326576,0.11861766,0.0025137055,0.008097558,0.037362184,0.02282065,0.67093295],"study_design_scores_gemma":[0.000104347455,0.0048091724,0.4936063,0.0061013387,0.0005558442,0.0014166718,0.17227075,0.008734344,0.015848352,0.023567507,0.2727891,0.00019626398],"about_ca_topic_score_codex":0.00095905125,"about_ca_topic_score_gemma":0.00089951075,"teacher_disagreement_score":0.0071844053,"about_ca_system_score_codex":0.0020780703,"about_ca_system_score_gemma":0.0022862244,"threshold_uncertainty_score":0.03799522},"labels":[],"label_agreement":null},{"id":"W2088108067","doi":"10.1016/j.scico.2013.11.025","title":"Understanding database schema evolution: A case study","year":2013,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science; Fonds De La Recherche Scientifique - FNRS","keywords":"Computer science; Database schema; Database; Reverse engineering; Schema (genetic algorithms); Database design; Schema migration; Database model; Information schema; Data science; Schema evolution; Software engineering; Information retrieval; Semi-structured model; Programming language","score_opus":0.08599344345401087,"score_gpt":0.31329145225617433,"score_spread":0.22729800880216344,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088108067","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88773364,0.00073460775,0.09236269,0.00398237,0.00004850999,0.0003514585,0.0005608238,0.0006450215,0.0135809295],"genre_scores_gemma":[0.9003116,0.00050922856,0.09400111,0.00034013807,0.000017498965,0.000064294996,0.0005756981,0.00017219991,0.0040083104],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9972711,0.001325224,0.0001689479,0.0004135736,0.0006640574,0.00015705696],"domain_scores_gemma":[0.9720416,0.022020843,0.0011820372,0.0024351797,0.0018933384,0.00042704612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003870698,0.00040695933,0.00031070734,0.001260825,0.0016444268,0.0032413402,0.0021028677,0.0034028422,0.0022961362],"category_scores_gemma":[0.02465983,0.0004868298,0.0005201671,0.0017928205,0.0013036248,0.0054269484,0.0018178127,0.0021819335,0.00033324352],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010744699,0.0036392743,0.12176963,0.0020661461,0.00020064572,0.026919477,0.16972168,0.030172082,0.032813102,0.076967366,0.0124843735,0.5221718],"study_design_scores_gemma":[0.00033861073,0.0015217784,0.05660826,0.0011896908,0.0005790081,0.042197675,0.15060453,0.38237852,0.088960916,0.07943961,0.19586419,0.00031709223],"about_ca_topic_score_codex":0.007293737,"about_ca_topic_score_gemma":0.014107311,"teacher_disagreement_score":0.007293737,"about_ca_system_score_codex":0.0016750466,"about_ca_system_score_gemma":0.0015050059,"threshold_uncertainty_score":0.02047044},"labels":[],"label_agreement":null},{"id":"W2088481278","doi":"10.5555/2486788.2487071","title":"4th international workshop on managing technical debt (MTD 2013)","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Credibility; Debt; Software; Computer science; Engineering management; Software engineering; Software development; Engineering; Systems engineering; Engineering ethics; Business; Political science; Finance","score_opus":0.028552896519075485,"score_gpt":0.2845764858648499,"score_spread":0.2560235893457744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088481278","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026835298,0.040879395,0.43463096,0.112771206,0.09045027,0.0016884893,0.0049228063,0.008467396,0.27935418],"genre_scores_gemma":[0.15433326,0.02138884,0.26694095,0.015732145,0.01455356,0.0021030984,0.014863307,0.0046991226,0.50538576],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9933515,0.0026727964,0.00040262894,0.0008913575,0.0017964392,0.0008853264],"domain_scores_gemma":[0.9923052,0.0018267451,0.00034921293,0.0011668835,0.0023611088,0.001990878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009968475,0.0014938605,0.001089388,0.0020612609,0.002676835,0.008983675,0.003317899,0.0047813063,0.057142675],"category_scores_gemma":[0.013324753,0.0007845201,0.0019579309,0.0015592038,0.0013059382,0.0076043033,0.010540503,0.0054931934,0.019823095],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032850582,0.00042200356,0.0010664061,0.0005698227,0.000064380205,0.0005156631,0.0026444409,0.0021727865,0.0045450707,0.031135196,0.62702006,0.32951555],"study_design_scores_gemma":[0.000032871092,0.00008971784,0.00074513105,0.0004492374,0.000028280312,0.0003166476,0.001250046,0.0026724406,0.0018526927,0.014773166,0.9777495,0.00004016465],"about_ca_topic_score_codex":0.0032523132,"about_ca_topic_score_gemma":0.005978777,"teacher_disagreement_score":0.057142675,"about_ca_system_score_codex":0.0026765245,"about_ca_system_score_gemma":0.004798299,"threshold_uncertainty_score":0.19116127},"labels":[],"label_agreement":null},{"id":"W2088569554","doi":"10.1109/mtd.2014.17","title":"Managing Technical Debt in Database Schemas of Critical Software","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Technical debt; Computer science; Database; Context (archaeology); Database schema; Software quality; Software engineering; Software; Software development; Database design; Data science; Risk analysis (engineering); Programming language","score_opus":0.018771384330771797,"score_gpt":0.29926621445905544,"score_spread":0.28049483012828363,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088569554","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6461692,0.0018949449,0.32987374,0.0035515567,0.00014603935,0.00054140907,0.0004306617,0.0017318911,0.015660571],"genre_scores_gemma":[0.8756278,0.0007327909,0.11764609,0.0006011487,0.00007167427,0.00020132512,0.00067748554,0.0004948728,0.0039466764],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99110514,0.0031310369,0.0011079792,0.0009767126,0.0030327535,0.00064627297],"domain_scores_gemma":[0.9502364,0.013677355,0.010640943,0.015541773,0.008065196,0.0018383622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013940417,0.0003552511,0.00055560993,0.0020850066,0.0018162957,0.007289406,0.0034332534,0.0019074642,0.001773286],"category_scores_gemma":[0.05389061,0.0010341977,0.00045610542,0.0029757456,0.0030625241,0.010559558,0.0056524905,0.0020325168,0.00046427918],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006901857,0.00042918033,0.10058158,0.0010240986,0.00020453084,0.0113513125,0.032695204,0.05373235,0.036047474,0.54589206,0.009380823,0.20797123],"study_design_scores_gemma":[0.00025765414,0.00082290254,0.03332483,0.00088792515,0.00027834516,0.0155161135,0.02094135,0.2630626,0.04421344,0.3977038,0.22268985,0.00030123166],"about_ca_topic_score_codex":0.002762351,"about_ca_topic_score_gemma":0.002266414,"teacher_disagreement_score":0.013940417,"about_ca_system_score_codex":0.0030748404,"about_ca_system_score_gemma":0.002825332,"threshold_uncertainty_score":0.073724866},"labels":[],"label_agreement":null},{"id":"W2088585190","doi":"10.1145/1806916.1806918","title":"Modeling web quality using a probabilistic approach","year":2010,"lang":"en","type":"article","venue":"ACM Transactions on the Web","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Université de Montréal","funders":"","keywords":"Navigability; Computer science; Quality (philosophy); Task (project management); Probabilistic logic; Process (computing); Bayesian network; Data mining; Statistical model; Machine learning; Artificial intelligence; Systems engineering; Engineering","score_opus":0.07936546576008312,"score_gpt":0.31188893004664103,"score_spread":0.2325234642865579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088585190","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051837992,0.00014780344,0.94533616,0.0002331554,0.000009604357,0.0001084411,0.00011765468,0.000217754,0.0019914135],"genre_scores_gemma":[0.7966071,0.00041705277,0.20057341,0.0000631052,0.000052658175,0.0005007111,0.00026307354,0.000058552654,0.0014642534],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959734,0.0020214173,0.00019392275,0.00049703725,0.001140081,0.00017422123],"domain_scores_gemma":[0.98004824,0.01463721,0.0025666754,0.00085913937,0.0016447274,0.00024402363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052668024,0.0010176648,0.0007221107,0.003446608,0.00052081083,0.0023480756,0.0015483157,0.0013881306,0.0016662492],"category_scores_gemma":[0.02998659,0.00082152715,0.0012503786,0.002130435,0.001154682,0.0033723502,0.0012964685,0.0012444692,0.0003037294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012533815,0.00024046442,0.015789224,0.0001444803,0.00020220439,0.00014029176,0.0005430499,0.9015982,0.0017696426,0.036504798,0.00041685553,0.04252541],"study_design_scores_gemma":[0.000011827484,0.00004360175,0.0022850453,0.0000117761165,0.00003591593,0.00003559136,0.000037051115,0.983705,0.00017523895,0.013382238,0.00025815732,0.000018541756],"about_ca_topic_score_codex":0.008717913,"about_ca_topic_score_gemma":0.0072928043,"teacher_disagreement_score":0.008717913,"about_ca_system_score_codex":0.001364486,"about_ca_system_score_gemma":0.00094984897,"threshold_uncertainty_score":0.027853847},"labels":[],"label_agreement":null},{"id":"W2088785743","doi":"10.1109/shark-adi.2007.12","title":"Social Factors Relevant to Capturing Design Decisions","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Casual; Computer science; Participatory design; Conversation; Software design; Management science; Software; Human–computer interaction; Knowledge management; Software development; Engineering; Psychology; Operations management","score_opus":0.08065930251151042,"score_gpt":0.3297008610409194,"score_spread":0.24904155852940896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088785743","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8469432,0.00056673586,0.0969232,0.002623965,0.00010694995,0.0012961233,0.00041178922,0.00017232646,0.050955776],"genre_scores_gemma":[0.98404944,0.00012734398,0.014130499,0.00011301154,0.000012897359,0.0004906466,0.0000835786,0.000030072086,0.0009624936],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9291035,0.05408154,0.0035720174,0.0024403709,0.008800451,0.0020019934],"domain_scores_gemma":[0.73050594,0.22110507,0.0210348,0.0090570925,0.016145062,0.00215201],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03205023,0.0008247485,0.00037112655,0.0034538337,0.0046053166,0.0057917396,0.0011470437,0.0011147052,0.0035569107],"category_scores_gemma":[0.15146571,0.0005821979,0.00047304263,0.0027548661,0.0057498347,0.0049800337,0.0034472488,0.0018267723,0.0004640122],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031423854,0.000233193,0.09621827,0.0014622567,0.00011244088,0.0010677815,0.75853837,0.0016755848,0.011409009,0.03850176,0.0014235137,0.08904357],"study_design_scores_gemma":[0.000089213805,0.00035451943,0.08268623,0.0008731929,0.0001592261,0.0008964159,0.7481199,0.008372131,0.012459124,0.058203287,0.08753145,0.00025532505],"about_ca_topic_score_codex":0.004778813,"about_ca_topic_score_gemma":0.0069770706,"teacher_disagreement_score":0.03205023,"about_ca_system_score_codex":0.0047103013,"about_ca_system_score_gemma":0.0045756297,"threshold_uncertainty_score":0.16949987},"labels":[],"label_agreement":null},{"id":"W2089736546","doi":"10.1145/1117696.1117718","title":"Study of novice programmers using Eclipse and Gild","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Eclipse; Computer science; Programming language; Astronomy; Physics","score_opus":0.04235982628760339,"score_gpt":0.31986306750956633,"score_spread":0.27750324122196296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089736546","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9978058,0.0000989361,0.000951319,0.00006310474,0.000008118583,0.0000913163,0.00003322628,0.00002517896,0.00092308165],"genre_scores_gemma":[0.9912373,0.00030314445,0.0045684027,0.0003588635,0.000021947002,0.00020056813,0.00022972935,0.000059752656,0.003020359],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99659985,0.0014452911,0.00027568376,0.00048098539,0.000799083,0.00039914658],"domain_scores_gemma":[0.96445704,0.02405903,0.0014961545,0.0016548227,0.0047891554,0.0035438656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048210183,0.00077361346,0.0010172229,0.0015977435,0.0018309479,0.0016510835,0.0011474956,0.0012484014,0.0019310969],"category_scores_gemma":[0.025309157,0.00068045635,0.0005155853,0.00095399923,0.00096844777,0.0015246891,0.0018942723,0.0019069749,0.00077742676],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015967091,0.016890679,0.29134023,0.0018794902,0.0001528146,0.009506163,0.44751748,0.0013846662,0.041930918,0.00081626896,0.003993071,0.18299149],"study_design_scores_gemma":[0.00063682214,0.034747567,0.6691337,0.00046278088,0.00019762441,0.009510283,0.22962096,0.0062049916,0.019642508,0.0007471972,0.028693689,0.00040174558],"about_ca_topic_score_codex":0.0024711182,"about_ca_topic_score_gemma":0.0046509714,"teacher_disagreement_score":0.0048210183,"about_ca_system_score_codex":0.0005828936,"about_ca_system_score_gemma":0.0006692471,"threshold_uncertainty_score":0.025496304},"labels":[],"label_agreement":null},{"id":"W2089779090","doi":"10.1109/icsm.2007.4362639","title":"How Programmers can Turn Comments into Waypoints for Code Navigation","year":2007,"lang":"en","type":"article","venue":"Proceedings/Proceedings - Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; University of Victoria","funders":"","keywords":"Computer science; Software engineering; Software; Software development; Code (set theory); Software maintenance; Human–computer interaction; Programming language","score_opus":0.03317168845903403,"score_gpt":0.2878660248895109,"score_spread":0.2546943364304769,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089779090","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10137955,0.0002908399,0.8465379,0.0024688314,0.00024587926,0.000342248,0.00031511148,0.021394275,0.027025338],"genre_scores_gemma":[0.44385836,0.0003777416,0.5322893,0.0004947847,0.00005460805,0.00024516074,0.00059791066,0.004528371,0.01755368],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9945062,0.0029532579,0.00033406727,0.0008131429,0.0010773713,0.00031588643],"domain_scores_gemma":[0.95021206,0.034016903,0.0024559163,0.008098454,0.0040826816,0.001134064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007799086,0.00089048425,0.00041003956,0.0019441157,0.001249159,0.005862056,0.0016660857,0.0021501225,0.008858351],"category_scores_gemma":[0.05932818,0.0011026621,0.00079267164,0.0013979119,0.0024791546,0.01201988,0.0041639362,0.0021085097,0.0045900955],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080229004,0.0004918958,0.017062254,0.0011056435,0.00009574024,0.0017852982,0.074738316,0.0069620265,0.05671021,0.11646546,0.032555446,0.6912254],"study_design_scores_gemma":[0.0004756165,0.00084362674,0.009764236,0.0012472966,0.00030915314,0.0035271596,0.029380677,0.08433679,0.10127959,0.17684296,0.59129184,0.0007010238],"about_ca_topic_score_codex":0.002507136,"about_ca_topic_score_gemma":0.0042043556,"teacher_disagreement_score":0.008858351,"about_ca_system_score_codex":0.0005077232,"about_ca_system_score_gemma":0.0014435257,"threshold_uncertainty_score":0.041245997},"labels":[],"label_agreement":null},{"id":"W2089862216","doi":"10.1109/qsic.2014.45","title":"Anti-pattern Mutations and Fault-proneness","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Commit; Computer science; Software evolution; Software quality; Software fault tolerance; Fault (geology); Software system; Software; Software development; Software construction; Programming language; Biology; Database","score_opus":0.013835058301845277,"score_gpt":0.26073561751367547,"score_spread":0.24690055921183018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089862216","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99838424,0.000047866408,0.0008514922,0.000042737694,0.0000015760847,0.0000071584486,0.000018124207,0.000012320892,0.0006345368],"genre_scores_gemma":[0.9993363,0.000020351112,0.00047259667,0.0000071688223,0.0000013675414,0.0000041722997,0.00001609388,0.0000039228034,0.00013812621],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975431,0.00067859184,0.00025207706,0.00034770896,0.0009911752,0.00018734438],"domain_scores_gemma":[0.909199,0.046765365,0.033298336,0.0039386787,0.0052423994,0.0015562609],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018489276,0.00024892404,0.00017591356,0.0014345093,0.00037551054,0.00079635275,0.00042284856,0.0005383454,0.0010955371],"category_scores_gemma":[0.034593295,0.00017687108,0.00020335737,0.0007282062,0.0010400317,0.001102877,0.0007554301,0.0006677821,0.0001182931],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001653617,0.00019136007,0.955028,0.00007405857,0.000073283554,0.00032537855,0.0024228762,0.0025022053,0.007842513,0.000991778,0.00010627193,0.0302769],"study_design_scores_gemma":[0.000008092705,0.00016167761,0.9894162,0.000016392692,0.000025777152,0.0005787292,0.00080287474,0.004432158,0.0027201625,0.0014112535,0.00040879968,0.000017885459],"about_ca_topic_score_codex":0.0012574028,"about_ca_topic_score_gemma":0.0020089229,"teacher_disagreement_score":0.0018489276,"about_ca_system_score_codex":0.0005810429,"about_ca_system_score_gemma":0.00046990326,"threshold_uncertainty_score":0.009778142},"labels":[],"label_agreement":null},{"id":"W2089863833","doi":"10.1109/seaa.2011.59","title":"Empirical Evaluation of Mixed-Project Defect Prediction Models","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Metric (unit); Predictive modelling; Product metric; Data modeling; Data mining; Focus (optics); Project management; Data collection; Software bug; Empirical research; Data science; Machine learning; Database; Software; Engineering; Systems engineering; Statistics","score_opus":0.232463589436113,"score_gpt":0.36341172342659717,"score_spread":0.13094813399048416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089863833","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9520673,0.0014229808,0.040929973,0.0005728811,0.00007773569,0.00011360507,0.0020909642,0.0014214753,0.0013030377],"genre_scores_gemma":[0.9751549,0.00017753239,0.019986654,0.000059254744,0.00003466661,0.00009840259,0.003972834,0.00007983352,0.0004359675],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9904614,0.0058356654,0.0006972085,0.0016638326,0.0010226511,0.0003192087],"domain_scores_gemma":[0.9065637,0.07468996,0.0036126457,0.0068992423,0.0064919367,0.0017424239],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027553266,0.002420178,0.0018115704,0.0028137232,0.0006277511,0.0021483605,0.0034485136,0.0019637987,0.0013179227],"category_scores_gemma":[0.048232336,0.0007927321,0.0012348837,0.001942731,0.00092902425,0.0035662486,0.0019205742,0.0022152627,0.0008474067],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019521342,0.0012551018,0.15952428,0.00033617183,0.0010118156,0.00021826233,0.00026089142,0.74680185,0.00089610525,0.0010337569,0.004418963,0.082290664],"study_design_scores_gemma":[0.000037718208,0.00025958364,0.0055086245,0.00001715973,0.000038797083,0.000046199522,0.00004843696,0.9928939,0.00038144572,0.0005631766,0.00018957235,0.000015344594],"about_ca_topic_score_codex":0.008397335,"about_ca_topic_score_gemma":0.0062645716,"teacher_disagreement_score":0.027553266,"about_ca_system_score_codex":0.0014569971,"about_ca_system_score_gemma":0.0011015907,"threshold_uncertainty_score":0.14571732},"labels":[],"label_agreement":null},{"id":"W2089991372","doi":"10.1109/wcre.2007.32","title":"Lossless Comparison of Nested Software Decompositions","year":2007,"lang":"en","type":"article","venue":"Proceedings - Working Conference on Reverse Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Decomposition; Lossless compression; Software; Cluster analysis; Software system; Theoretical computer science; Nested set model; Algorithm; Data mining; Distributed computing; Programming language; Data compression; Relational database; Artificial intelligence","score_opus":0.04362021653374086,"score_gpt":0.3034449051996365,"score_spread":0.25982468866589564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089991372","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27199894,0.0004452937,0.7173106,0.0001031726,0.00012395017,0.00015873923,0.00047496095,0.0025804862,0.0068039],"genre_scores_gemma":[0.7360771,0.00013408555,0.25963864,0.00005829615,0.000026817595,0.00012270786,0.0011473962,0.00045886438,0.0023361277],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99470633,0.0014514695,0.00033278042,0.0006242297,0.0024922653,0.00039301417],"domain_scores_gemma":[0.9795919,0.009475248,0.0012940598,0.0040103267,0.0050934698,0.0005349843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059891576,0.0007670749,0.00076248165,0.0031079163,0.0006692752,0.0017943231,0.0013263342,0.00055690383,0.0035470815],"category_scores_gemma":[0.034361616,0.00032739414,0.0005265499,0.0016169362,0.0008636284,0.0036426212,0.002532419,0.0008125327,0.00071980286],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039663156,0.00037733518,0.01674699,0.0005385838,0.0002157809,0.00035907506,0.0008494918,0.14847742,0.057945214,0.06436729,0.0049458332,0.7012106],"study_design_scores_gemma":[0.00012481352,0.00090508745,0.009944542,0.00007527187,0.000101057994,0.00050537,0.00050755945,0.8512016,0.066477135,0.06248845,0.0075920094,0.00007710835],"about_ca_topic_score_codex":0.0010692292,"about_ca_topic_score_gemma":0.0015019036,"teacher_disagreement_score":0.0059891576,"about_ca_system_score_codex":0.0009507332,"about_ca_system_score_gemma":0.0008844787,"threshold_uncertainty_score":0.031674087},"labels":[],"label_agreement":null},{"id":"W2090222336","doi":"10.1145/1985404.1985411","title":"Automated type-3 clone oracle using Levenshtein metric","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Oracle; clone (Java method); Computer science; Construct (python library); Metric (unit); Scalability; Relation (database); Data mining; Programming language; Theoretical computer science; Database; Engineering","score_opus":0.0943181722566206,"score_gpt":0.30913082433026706,"score_spread":0.21481265207364647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2090222336","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11898165,0.00046260536,0.87029237,0.00012429274,0.000031738644,0.00018584779,0.000353845,0.0074133635,0.0021541943],"genre_scores_gemma":[0.6003339,0.00018430209,0.39634955,0.000047918413,0.00002667733,0.00014835848,0.00096175226,0.00048129226,0.0014662622],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9903413,0.0016373303,0.0013578127,0.001269214,0.0048501533,0.00054415496],"domain_scores_gemma":[0.9703856,0.011538849,0.0037669246,0.0053999373,0.0081989495,0.000709613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044328035,0.0007950084,0.0014785047,0.004825802,0.0009189095,0.0036976298,0.0016601146,0.0011885093,0.0014048404],"category_scores_gemma":[0.029684188,0.00037911197,0.0008718316,0.0025575066,0.0013077386,0.0041242978,0.0020676109,0.00095297414,0.0007052392],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007977334,0.00018900291,0.07137206,0.0005514882,0.00016219135,0.0006397607,0.0017615438,0.046574652,0.085862316,0.036858976,0.0038836822,0.75134665],"study_design_scores_gemma":[0.000049963706,0.0008095587,0.030018859,0.00009406799,0.00014558618,0.0020985443,0.00064500055,0.66376764,0.25273496,0.03388601,0.015511568,0.0002382876],"about_ca_topic_score_codex":0.0025637615,"about_ca_topic_score_gemma":0.0019756379,"teacher_disagreement_score":0.004825802,"about_ca_system_score_codex":0.0015523339,"about_ca_system_score_gemma":0.001392211,"threshold_uncertainty_score":0.023443162},"labels":[],"label_agreement":null},{"id":"W2090398160","doi":"10.1109/esem.2011.22","title":"How Good is Your Comment? A Study of Comments in Java Programs","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Java; Structuring; Code review; Focus (optics); Code (set theory); Taxonomy (biology); Point (geometry); Empirical research; Source code; Software engineering; Static program analysis; Open source; Data science; Programming language; World Wide Web; Software development; Software; Set (abstract data type)","score_opus":0.1038988231966806,"score_gpt":0.30234095690741153,"score_spread":0.19844213371073094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2090398160","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99643457,0.00018250558,0.0007196262,0.0006182469,0.00001900981,0.000060205024,0.00005330376,0.000012589475,0.001899959],"genre_scores_gemma":[0.998156,0.00022386892,0.00053541403,0.00025833343,0.000027970475,0.000090681046,0.00006387892,0.000023224165,0.0006204894],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.98103505,0.0129216695,0.0012881429,0.001177586,0.0028412414,0.00073622924],"domain_scores_gemma":[0.5935417,0.33747908,0.035177045,0.0037939232,0.025677416,0.004330799],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016748017,0.0003952461,0.00053196,0.002478777,0.0034581712,0.0035388446,0.0010012095,0.0019578729,0.0013567253],"category_scores_gemma":[0.15374291,0.00043910416,0.0003490404,0.002769634,0.0037958925,0.005076062,0.0022263953,0.002419867,0.00039374953],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036047355,0.0006487157,0.17247948,0.00071305735,0.000050646562,0.0008363807,0.79145914,0.0001056949,0.0025131295,0.0010390102,0.0016819915,0.028112223],"study_design_scores_gemma":[0.000034859557,0.0010323017,0.28503278,0.00054442487,0.00004957983,0.00071555795,0.70033836,0.000865075,0.0013232755,0.0006221125,0.009335889,0.00010573007],"about_ca_topic_score_codex":0.0032140804,"about_ca_topic_score_gemma":0.004619084,"teacher_disagreement_score":0.016748017,"about_ca_system_score_codex":0.0021118321,"about_ca_system_score_gemma":0.001329127,"threshold_uncertainty_score":0.08857298},"labels":[],"label_agreement":null},{"id":"W2090432523","doi":"10.1007/s10664-008-9076-6","title":"“Cloning considered harmful” considered harmful: patterns of cloning in software","year":2008,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":363,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cloning (programming); Source code; Computer science; Software engineering; Code (set theory); Software system; Software; clone (Java method); Maintainability; Web application; Codebase; Programming language; World Wide Web; Biology; Genetics","score_opus":0.044662327963924,"score_gpt":0.2815639292607685,"score_spread":0.2369016012968445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2090432523","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9471351,0.00061161804,0.02840479,0.0043823253,0.000050760143,0.000067672154,0.00023786865,0.00016148841,0.01894835],"genre_scores_gemma":[0.9948925,0.000085587446,0.0040600584,0.00030649663,0.000016922082,0.000020732556,0.00008091426,0.000048955255,0.00048778555],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98219854,0.009560371,0.0017342173,0.0014244842,0.0036944733,0.0013879667],"domain_scores_gemma":[0.7696085,0.1448309,0.041810166,0.029270563,0.011050499,0.003429307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012257816,0.00030273636,0.00054259703,0.003719789,0.0043408195,0.0040287087,0.0013750414,0.0023844843,0.003393452],"category_scores_gemma":[0.1295866,0.00049913855,0.00051375886,0.00789833,0.0094178235,0.0082446635,0.0051291357,0.0038294152,0.0003049181],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073648285,0.00022182506,0.56268966,0.00059916044,0.00024311933,0.0007302896,0.11433606,0.0015375875,0.0056407657,0.20881662,0.0054363683,0.099012084],"study_design_scores_gemma":[0.00008643669,0.00033463247,0.48848248,0.00089652865,0.00042711498,0.0044539226,0.10480571,0.008696132,0.0077897776,0.34992328,0.033862174,0.00024180974],"about_ca_topic_score_codex":0.0051151607,"about_ca_topic_score_gemma":0.0062507805,"teacher_disagreement_score":0.012257816,"about_ca_system_score_codex":0.0019879965,"about_ca_system_score_gemma":0.0034266147,"threshold_uncertainty_score":0.06482631},"labels":[],"label_agreement":null},{"id":"W2090907135","doi":"10.1145/2568225.2568268","title":"Understanding JavaScript event-based interactions","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Intel Corporation","keywords":"Computer science; JavaScript; Unobtrusive JavaScript; Program comprehension; Asynchronous communication; Event (particle physics); Web application; Task (project management); Visualization; Software engineering; World Wide Web; Programming language; Software; Human–computer interaction; Software system; Artificial intelligence; Rich Internet application","score_opus":0.09474525985388557,"score_gpt":0.30118461458695595,"score_spread":0.20643935473307037,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2090907135","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15110949,0.00023518139,0.8354155,0.00018684908,0.000028791525,0.00014382247,0.0004923898,0.0075889146,0.004799155],"genre_scores_gemma":[0.6762707,0.00037358483,0.31891927,0.000095141964,0.000024271232,0.00011589464,0.0012853822,0.00081799173,0.0020977217],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9992888,0.00017042574,0.00004801045,0.00018737359,0.0002518758,0.000053484444],"domain_scores_gemma":[0.99733573,0.0017706001,0.00031321065,0.00022533679,0.0002702884,0.00008495234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006615064,0.00077230815,0.0003769716,0.00078751997,0.00035707685,0.0019120415,0.0010224028,0.0009716109,0.0013000217],"category_scores_gemma":[0.0051177456,0.00045280572,0.0004266418,0.00045331637,0.0004301653,0.0020874995,0.0009297388,0.0011207963,0.0005562561],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080385804,0.00071554794,0.038688064,0.0013328784,0.00015309262,0.0022746185,0.010105019,0.19354339,0.25444004,0.02671395,0.0077031143,0.4635264],"study_design_scores_gemma":[0.000032687378,0.00008626921,0.015815489,0.00006822452,0.000042610267,0.00042451534,0.00050340383,0.9192127,0.03529034,0.014576846,0.013902482,0.000044418644],"about_ca_topic_score_codex":0.005366642,"about_ca_topic_score_gemma":0.0051406086,"teacher_disagreement_score":0.005366642,"about_ca_system_score_codex":0.00051406893,"about_ca_system_score_gemma":0.00061358564,"threshold_uncertainty_score":0.010670841},"labels":[],"label_agreement":null},{"id":"W2090908516","doi":"10.1145/2430536.2430539","title":"Facilitating the transition from use case models to analysis models","year":2013,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":114,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Pearl Therapeutics; Fonds National de la Recherche Luxembourg","keywords":"Use Case Diagram; Computer science; Ambiguity; Set (abstract data type); Unified Modeling Language; Sequence diagram; Class diagram; Class (philosophy); Quality (philosophy); Data mining; Natural language processing; Artificial intelligence; Programming language; Software","score_opus":0.15956612434369744,"score_gpt":0.31830746037286617,"score_spread":0.15874133602916873,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2090908516","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01792267,0.0001074642,0.9701198,0.000608583,0.00004599776,0.0016226632,0.00025636828,0.005306606,0.0040098275],"genre_scores_gemma":[0.07200021,0.00017642182,0.9220719,0.0002585861,0.000030271516,0.001715,0.00083598925,0.001229864,0.001681772],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9399031,0.03390531,0.005365497,0.0056413217,0.013567452,0.0016173404],"domain_scores_gemma":[0.8098975,0.12667574,0.006813629,0.04305445,0.011898023,0.0016607273],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.036810204,0.0020029242,0.0013586127,0.0039366577,0.0013706106,0.010157862,0.005424087,0.0030195853,0.006089839],"category_scores_gemma":[0.14628083,0.003400991,0.0025945222,0.0022431528,0.0022597301,0.013599474,0.011685369,0.0070995544,0.004046622],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008878723,0.0022230248,0.009445328,0.0021565855,0.00030909045,0.0035932434,0.046303824,0.05535394,0.062059544,0.20462935,0.013731982,0.5993062],"study_design_scores_gemma":[0.00037055844,0.0008276296,0.0044268463,0.0017873782,0.0003050089,0.0015647135,0.006243342,0.51316196,0.077855214,0.1715712,0.22136275,0.0005233063],"about_ca_topic_score_codex":0.0024021138,"about_ca_topic_score_gemma":0.0025722855,"teacher_disagreement_score":0.036810204,"about_ca_system_score_codex":0.0025842108,"about_ca_system_score_gemma":0.0054863286,"threshold_uncertainty_score":0.1946733},"labels":[],"label_agreement":null},{"id":"W2090915952","doi":"10.1080/10429247.2007.11431746","title":"A Project Management Approach to Using Simulation for Cost Estimation on Large, Complex Software Development Projects","year":2007,"lang":"en","type":"article","venue":"Engineering Management Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Saskatchewan; Portland State University","keywords":"Software project management; Project management; Software development; Computer science; Cost estimate; Software; Software sizing; Systems engineering; Software development process; Estimation; Verification and validation; Software construction; Software engineering; Industrial engineering; Engineering; Operations management","score_opus":0.07778606872762694,"score_gpt":0.33518546918437053,"score_spread":0.2573994004567436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2090915952","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070992135,0.00021909742,0.9841218,0.001198334,0.00006137661,0.00013951086,0.000050782648,0.00034000265,0.0067699775],"genre_scores_gemma":[0.20256317,0.00094173604,0.79251903,0.00024999323,0.00008778845,0.00065074576,0.0001214596,0.000118605945,0.0027473376],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9931496,0.0046975445,0.0002644282,0.00034247825,0.0013898938,0.00015607366],"domain_scores_gemma":[0.98984563,0.006404709,0.00077822106,0.0010089839,0.0016362953,0.00032620676],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006407044,0.0012081198,0.00076483004,0.0027804133,0.0014132169,0.0040643876,0.0021045774,0.0015582822,0.0039535365],"category_scores_gemma":[0.01583928,0.0009341135,0.001143607,0.0029822362,0.0014304973,0.0035366123,0.0019452568,0.0021730196,0.0005585984],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008690866,0.0002386261,0.0035595135,0.00015046142,0.00017712395,0.00017903026,0.00052100903,0.6789128,0.0019933556,0.16103753,0.003801591,0.14934209],"study_design_scores_gemma":[0.000050782084,0.00017548856,0.00064112194,0.00008730361,0.000057326844,0.00012333335,0.00021810058,0.8979405,0.001406401,0.08190455,0.017323343,0.00007174622],"about_ca_topic_score_codex":0.0050999224,"about_ca_topic_score_gemma":0.007000029,"teacher_disagreement_score":0.006407044,"about_ca_system_score_codex":0.0023088218,"about_ca_system_score_gemma":0.0041341074,"threshold_uncertainty_score":0.03388405},"labels":[],"label_agreement":null},{"id":"W2091122638","doi":"10.1109/apsec.2011.51","title":"A Case Study of Measuring Degeneration of Software Architectures from a Defect Perspective","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Program for New Century Excellent Talents in University","keywords":"Computer science; Perspective (graphical); Compiler; Architecture; Degeneration (medical); Architectural pattern; Software architecture; Software engineering; Software; Software system; Programming language; Artificial intelligence; Software construction","score_opus":0.07015871070585016,"score_gpt":0.2706566237819026,"score_spread":0.20049791307605241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091122638","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98826885,0.00017599623,0.010711833,0.00006699515,0.000008431239,0.000055466724,0.000079405225,0.00008303616,0.00054989906],"genre_scores_gemma":[0.9855076,0.000070054026,0.014008214,0.000015460379,0.0000058489222,0.000021113701,0.00009180684,0.000019486351,0.00026050032],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99621564,0.0014472189,0.0003000159,0.00046971173,0.0013629715,0.00020445342],"domain_scores_gemma":[0.97497976,0.014589173,0.0029751281,0.0028666942,0.003900958,0.0006882954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030261644,0.00054713653,0.0006021099,0.0022750641,0.0008791396,0.0008347744,0.0010294038,0.0011561734,0.00033502132],"category_scores_gemma":[0.014847466,0.00028631746,0.00049014133,0.002376469,0.0012444716,0.0011315645,0.00082322996,0.0007851536,0.00009888174],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001765797,0.0018376975,0.46488148,0.0012420445,0.00046982174,0.0153262885,0.014812456,0.12015444,0.13383171,0.006404574,0.0021410857,0.23713255],"study_design_scores_gemma":[0.00021942497,0.00750836,0.45591986,0.00023531458,0.00041564592,0.020324172,0.00732355,0.3254441,0.16715032,0.0073404084,0.007836489,0.00028236915],"about_ca_topic_score_codex":0.0031032367,"about_ca_topic_score_gemma":0.0040332074,"teacher_disagreement_score":0.0031032367,"about_ca_system_score_codex":0.0009512346,"about_ca_system_score_gemma":0.00038073928,"threshold_uncertainty_score":0.016004026},"labels":[],"label_agreement":null},{"id":"W2091156716","doi":"10.1016/j.infsof.2009.08.001","title":"Software development effort prediction: A study on the factors impacting the accuracy of fuzzy logic systems","year":2009,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"King Fahd University of Petroleum and Minerals","keywords":"Fuzzy logic; Computer science; Robustness (evolution); Fuzzy electronics; Software; Machine learning; Software development process; Software development; Data mining; Fuzzy set; Artificial intelligence; Fuzzy set operations","score_opus":0.024011632092427623,"score_gpt":0.26694174166021956,"score_spread":0.24293010956779193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091156716","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9970445,0.0001345234,0.0020240091,0.00006370837,0.000003535695,0.000010336674,0.000043625238,0.000013974347,0.0006618396],"genre_scores_gemma":[0.9987931,0.000050395087,0.0009850581,0.000007723297,0.0000037083016,0.000003731033,0.000053856333,0.0000034201203,0.00009889206],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971233,0.0011383307,0.00025030764,0.00027090195,0.0010362681,0.00018089071],"domain_scores_gemma":[0.7790546,0.20319882,0.0074334205,0.002470054,0.007228549,0.0006145177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071905702,0.00032133542,0.00031485464,0.0012941615,0.00037917757,0.0010477883,0.00060396316,0.0006527975,0.00069894077],"category_scores_gemma":[0.07911137,0.00019950986,0.0005033531,0.0015847989,0.00032794802,0.0017168884,0.0002752841,0.0007450013,0.00014670283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00109698,0.0006164063,0.9013789,0.00012630972,0.00017140369,0.00016355563,0.0010933719,0.020657081,0.0020315973,0.0005566409,0.00034210703,0.07176565],"study_design_scores_gemma":[0.00005151242,0.0014325918,0.734272,0.00009264449,0.00026757075,0.00028060516,0.0012805043,0.2529411,0.00736925,0.0012988521,0.0006597182,0.00005361729],"about_ca_topic_score_codex":0.00941362,"about_ca_topic_score_gemma":0.0050953375,"teacher_disagreement_score":0.00941362,"about_ca_system_score_codex":0.0008183367,"about_ca_system_score_gemma":0.00076781644,"threshold_uncertainty_score":0.038027823},"labels":[],"label_agreement":null},{"id":"W2091379928","doi":"10.1145/774833.774839","title":"EVolve","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Visualization; Java; Extensibility; Protocol (science); Context (archaeology); Data visualization; Variety (cybernetics); Human–computer interaction; Programming language; Software engineering; Data mining; Artificial intelligence","score_opus":0.016618880582050233,"score_gpt":0.2543735368530366,"score_spread":0.23775465627098635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091379928","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017399417,0.0011070297,0.7146013,0.0010647556,0.00039589193,0.00029145702,0.012186216,0.1684234,0.08453059],"genre_scores_gemma":[0.14898045,0.0023122767,0.68737763,0.00092900184,0.00015835387,0.00086496834,0.03514088,0.050881796,0.07335466],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994937,0.00009945226,0.00003928186,0.0001125685,0.0002050151,0.000049887178],"domain_scores_gemma":[0.9991799,0.00027068332,0.000052591393,0.00025454187,0.0001653233,0.00007691958],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00083614566,0.0008095156,0.0005000734,0.0010318878,0.0006931518,0.002167463,0.0012247786,0.0009369695,0.025643399],"category_scores_gemma":[0.0032885082,0.0005270866,0.0010118112,0.00086914713,0.00039817323,0.0031605,0.0021993343,0.0014969555,0.0076136263],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006958796,0.00017864627,0.0051268577,0.0012344026,0.00020038922,0.0005783675,0.0017053896,0.022471407,0.022827035,0.17537509,0.3044324,0.4651742],"study_design_scores_gemma":[0.00010905348,0.00008359251,0.002380774,0.00019895326,0.00007004734,0.0007525005,0.00023444972,0.07630874,0.022930171,0.05868006,0.83815,0.000101616155],"about_ca_topic_score_codex":0.0021187116,"about_ca_topic_score_gemma":0.0029241976,"teacher_disagreement_score":0.025643399,"about_ca_system_score_codex":0.0004201819,"about_ca_system_score_gemma":0.000805624,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2091483537","doi":"10.1155/2012/964064","title":"Evaluating the Effect of Control Flow on the Unit Testing Effort of Classes: An Empirical Analysis","year":2012,"lang":"en","type":"article","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Testability; Metric (unit); Computer science; Unit testing; Univariate; Java; Statistic; Software metric; Logistic regression; Regression testing; Reliability engineering; Data mining; Software quality; Software; Statistics; Machine learning; Programming language; Software system; Software development; Multivariate statistics; Mathematics; Operations management; Engineering","score_opus":0.0406293769444153,"score_gpt":0.3715752348543554,"score_spread":0.3309458579099401,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091483537","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9904772,0.00032220807,0.0081523,0.000073987234,0.000008550518,0.0000397117,0.00030477578,0.000094504314,0.0005268369],"genre_scores_gemma":[0.99790937,0.000054019587,0.0014960973,0.000016283368,0.000011256259,0.0000368688,0.0003481868,0.000023922883,0.00010402811],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97304326,0.015852502,0.0018237886,0.003190446,0.005199012,0.0008909798],"domain_scores_gemma":[0.2697136,0.68638986,0.02114067,0.014321235,0.0070168334,0.0014177867],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.034009825,0.0010143676,0.00075726554,0.0030000298,0.00031804715,0.0010532882,0.0014839154,0.001400425,0.0011916636],"category_scores_gemma":[0.23702161,0.00039298864,0.0013614952,0.0022411856,0.0017038082,0.0026104369,0.0009354772,0.0018706804,0.0003449885],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009831988,0.0006415708,0.9421676,0.000111345784,0.00068729284,0.00021395758,0.00031554335,0.029793128,0.0012066703,0.00034099768,0.00029278774,0.023245934],"study_design_scores_gemma":[0.000048278485,0.0021699457,0.8028958,0.000048178856,0.00033487438,0.00049019756,0.00017142149,0.18946931,0.0032683655,0.0005957536,0.000452888,0.000054977503],"about_ca_topic_score_codex":0.0019895365,"about_ca_topic_score_gemma":0.0014161537,"teacher_disagreement_score":0.9659902,"about_ca_system_score_codex":0.00075604045,"about_ca_system_score_gemma":0.00056598545,"threshold_uncertainty_score":0.17986333},"labels":[],"label_agreement":null},{"id":"W2091795249","doi":"10.4018/jssoe.2012070102","title":"Refactoring Flash Embedding Methods","year":2012,"lang":"en","type":"article","venue":"International Journal of Systems and Service-Oriented Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"JavaScript; Ajax; Computer science; Rich Internet application; Flash (photography); Unobtrusive JavaScript; Markup language; Embedding; Web application; HTML5; ActionScript; Programming language; World Wide Web; XML; Artificial intelligence","score_opus":0.021570531376186065,"score_gpt":0.3283223930711712,"score_spread":0.30675186169498514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091795249","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017469915,0.00049785513,0.92236173,0.00032564125,0.00035564307,0.00035587887,0.0005341058,0.05248033,0.00561883],"genre_scores_gemma":[0.11299839,0.00095952215,0.8293586,0.000569826,0.00012538533,0.00047862792,0.003901527,0.029940963,0.021667195],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99490213,0.00089189067,0.00073868403,0.0008554091,0.0023351463,0.0002766979],"domain_scores_gemma":[0.9775214,0.008030181,0.0010810009,0.007473147,0.0055514374,0.00034276038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005105747,0.001961937,0.0007464098,0.002071627,0.0006402178,0.0020206345,0.002882103,0.0012843491,0.0052356417],"category_scores_gemma":[0.021633673,0.001110369,0.0015482618,0.0010445416,0.0009433046,0.0034414812,0.0030947207,0.0023717077,0.003088115],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039290238,0.00034902297,0.0065514217,0.0018041678,0.00020799336,0.0017001758,0.004129611,0.008636223,0.06109337,0.023407882,0.033016752,0.8587104],"study_design_scores_gemma":[0.00024692138,0.00028722791,0.003202229,0.0010183896,0.0004091612,0.0026567192,0.0010390287,0.09431794,0.27901548,0.025737727,0.5916187,0.00045045308],"about_ca_topic_score_codex":0.0022938421,"about_ca_topic_score_gemma":0.0024934523,"teacher_disagreement_score":0.0052356417,"about_ca_system_score_codex":0.0007976981,"about_ca_system_score_gemma":0.0016888967,"threshold_uncertainty_score":0.027002096},"labels":[],"label_agreement":null},{"id":"W2091905052","doi":"10.1155/2010/307391","title":"A Quality Model for Conceptual Models of MDD Environments","year":2010,"lang":"en","type":"article","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Generalitat Valenciana; Ministerio de Economía y Competitividad","keywords":"Computer science; Conceptual model; Quality (philosophy); Software; Software engineering; Key (lock); Software quality; Systems engineering; Risk analysis (engineering); Data mining; Software development; Programming language; Engineering; Computer security; Database","score_opus":0.02494850797632375,"score_gpt":0.29624273413305474,"score_spread":0.271294226156731,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091905052","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031975606,0.00026097576,0.99088895,0.000516098,0.000048097638,0.00015808527,0.00019569535,0.00039105624,0.0043435157],"genre_scores_gemma":[0.13562733,0.00086963864,0.8568547,0.00028167936,0.00012099521,0.0008374483,0.00095956057,0.00042470606,0.0040239557],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9896171,0.0025384799,0.0012956518,0.0011392611,0.0048569427,0.0005525528],"domain_scores_gemma":[0.98016113,0.0075710784,0.0025999884,0.0038452956,0.005296495,0.00052587537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008470672,0.0012891833,0.00088873046,0.004127354,0.0010217885,0.0060849036,0.0031886932,0.003452826,0.0032935196],"category_scores_gemma":[0.028531447,0.0011719698,0.0026174171,0.0028411571,0.0028115632,0.0071440195,0.0033698974,0.0032661934,0.0012572274],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006010076,0.00007774116,0.0017529362,0.00026118138,0.000056392764,0.00038585195,0.00081449316,0.05829762,0.0023627013,0.8942717,0.0020497788,0.039609645],"study_design_scores_gemma":[0.00011162749,0.00019799262,0.0010071676,0.00035443695,0.0001668372,0.0009306546,0.0003244203,0.45725808,0.0038161927,0.44401062,0.09171922,0.00010274519],"about_ca_topic_score_codex":0.0068539977,"about_ca_topic_score_gemma":0.003976122,"teacher_disagreement_score":0.008470672,"about_ca_system_score_codex":0.0029400329,"about_ca_system_score_gemma":0.0027494323,"threshold_uncertainty_score":0.04479772},"labels":[],"label_agreement":null},{"id":"W2091911921","doi":"10.5555/1030818.1031030","title":"Construction engineering and project management III: monte carlo simulation for schedule risks","year":2003,"lang":"en","type":"article","venue":"Winter Simulation Conference","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Monte Carlo method; Schedule; Project management; Computer science; Probabilistic logic; Operations research; Systems engineering; Engineering; Mathematics; Statistics; Artificial intelligence","score_opus":0.06302238816012153,"score_gpt":0.3270304624963632,"score_spread":0.2640080743362417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091911921","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015663208,0.00041507976,0.9636865,0.00076085184,0.00006993691,0.000160584,0.00013978669,0.0004865476,0.01861751],"genre_scores_gemma":[0.5148233,0.0013345747,0.46187088,0.00026181943,0.00013380012,0.0013602291,0.00043596225,0.0007406981,0.019038735],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976913,0.00148392,0.00006159001,0.000104095874,0.00051096006,0.0001481712],"domain_scores_gemma":[0.9944988,0.0041546836,0.00031692628,0.00041384168,0.0004616667,0.0001541541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040547615,0.000685092,0.0009295585,0.0009336104,0.0006868541,0.0021226117,0.0012614003,0.0019416016,0.0058790226],"category_scores_gemma":[0.01198916,0.0009564606,0.00095101714,0.0015116682,0.0010413899,0.0013143903,0.0011048001,0.0018362463,0.0007362043],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000025268608,0.000019074763,0.00053305656,0.000019615001,0.000023764253,0.00002516052,0.00005009926,0.9630337,0.00019760052,0.027185008,0.0010655312,0.007822235],"study_design_scores_gemma":[0.00001567151,0.000017587234,0.00016350066,0.000017107252,0.0000075640946,0.0000134513075,0.000014348492,0.98392326,0.00019378688,0.012796896,0.0028275948,0.00000916557],"about_ca_topic_score_codex":0.016746877,"about_ca_topic_score_gemma":0.009664605,"teacher_disagreement_score":0.016746877,"about_ca_system_score_codex":0.0021640924,"about_ca_system_score_gemma":0.0034094916,"threshold_uncertainty_score":0.03329879},"labels":[],"label_agreement":null},{"id":"W2092358306","doi":"10.1145/2577554.2577559","title":"Genealogical insights into the facts and fictions of clone removal","year":2013,"lang":"en","type":"article","venue":"ACM SIGAPP Applied Computing Review","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Software maintenance; Empirical research; Computer science; Software; Data science; Software system; Biology; Programming language; Genetics; Statistics; Gene; Mathematics","score_opus":0.024851864374432757,"score_gpt":0.27188670534798987,"score_spread":0.2470348409735571,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2092358306","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60369813,0.030490132,0.20340756,0.06184863,0.0007725584,0.00020007744,0.0009527453,0.00027884395,0.09835127],"genre_scores_gemma":[0.95475423,0.012557219,0.023798142,0.0028825935,0.00040545018,0.000115007715,0.00029414214,0.00012343563,0.0050698677],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9938013,0.0022824793,0.0003892102,0.0012045034,0.0021142,0.000208332],"domain_scores_gemma":[0.87191266,0.1047196,0.0112113645,0.005860195,0.005624347,0.00067190186],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009247331,0.00036427513,0.0005047575,0.008124717,0.0032135595,0.0037826211,0.0015909077,0.0018082837,0.003376038],"category_scores_gemma":[0.0759509,0.00053063117,0.00030416987,0.0060887546,0.02460496,0.01639471,0.00251292,0.0033939031,0.00043781372],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021689014,0.00009934283,0.06266125,0.0010955818,0.00010573174,0.0021712473,0.088498235,0.004179469,0.0036113386,0.60992193,0.007611745,0.21982732],"study_design_scores_gemma":[0.00003413292,0.00021196171,0.06748318,0.0011634455,0.00010529334,0.005298147,0.04410802,0.0066431873,0.0035702835,0.66346604,0.20766953,0.00024675185],"about_ca_topic_score_codex":0.0036114375,"about_ca_topic_score_gemma":0.0049976823,"teacher_disagreement_score":0.009247331,"about_ca_system_score_codex":0.0033292342,"about_ca_system_score_gemma":0.001916158,"threshold_uncertainty_score":0.048905134},"labels":[],"label_agreement":null},{"id":"W2092731145","doi":"10.1145/1370750.1370787","title":"A newbie's guide to eclipse APIs","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Eclipse; Documentation; Computer science; Java; Software engineering; Popularity; Flexibility (engineering); Programming language","score_opus":0.023832653961991315,"score_gpt":0.2867507034800775,"score_spread":0.26291804951808617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2092731145","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014343844,0.010597787,0.522503,0.0040514097,0.0022084536,0.0021557605,0.046143938,0.1608436,0.25006163],"genre_scores_gemma":[0.0051485426,0.009213006,0.60579914,0.0033909485,0.0005854817,0.0030651297,0.06681876,0.0719987,0.23398034],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981535,0.00039915167,0.00033810592,0.00017864519,0.00079461373,0.00013597183],"domain_scores_gemma":[0.9939235,0.0027200351,0.00031355268,0.00093368447,0.001763326,0.00034603305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033452064,0.0022476942,0.0012268599,0.006968602,0.00094985217,0.0030629798,0.0032338877,0.0020128693,0.13087025],"category_scores_gemma":[0.013878049,0.003741,0.001361175,0.0044659944,0.00062007736,0.005932031,0.0024775972,0.004875855,0.16150124],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038443348,0.00007193762,0.0002299941,0.0005686161,0.000013409994,0.00017444245,0.0002086696,0.00037355698,0.0013667834,0.008982728,0.7953496,0.19262181],"study_design_scores_gemma":[0.000017633176,0.000009568951,0.0002634304,0.00018611534,0.0000044717062,0.00030199194,0.000021319222,0.0004887914,0.00029121627,0.0036888237,0.9947068,0.000019830202],"about_ca_topic_score_codex":0.0046224636,"about_ca_topic_score_gemma":0.012620408,"teacher_disagreement_score":0.13087025,"about_ca_system_score_codex":0.0009325729,"about_ca_system_score_gemma":0.0027278182,"threshold_uncertainty_score":0.43780464},"labels":[],"label_agreement":null},{"id":"W2092778012","doi":"10.1145/1188966.1188981","title":"All code coverage is not created equal","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Computer science; Schedule; Software engineering; Code coverage; Code (set theory); Static program analysis; Software; Software development; Programming language; Operating system; Set (abstract data type)","score_opus":0.02657410533940271,"score_gpt":0.27637849773905665,"score_spread":0.24980439239965394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2092778012","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18783075,0.0051495023,0.4008252,0.030659059,0.0018033778,0.0003279802,0.002099838,0.0066014375,0.3647028],"genre_scores_gemma":[0.82184845,0.0023000876,0.11128433,0.0047707898,0.0005715175,0.00039540997,0.001937195,0.0033224407,0.053569805],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98612,0.0031016532,0.0006339733,0.002400094,0.006449511,0.0012947777],"domain_scores_gemma":[0.95306087,0.017978081,0.0034499504,0.013724989,0.0099934405,0.0017926806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040715896,0.0008111742,0.001243127,0.0033901732,0.0025948714,0.0054555526,0.0017719411,0.0017392815,0.014500676],"category_scores_gemma":[0.045672167,0.0006540764,0.00088803546,0.003267165,0.005128707,0.009196869,0.0055164658,0.002180209,0.0047453507],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037305016,0.000105155814,0.014140168,0.0005677172,0.00025277684,0.0006168382,0.0023703477,0.007002096,0.011014979,0.40684053,0.03006368,0.52665263],"study_design_scores_gemma":[0.00006426944,0.00034902873,0.018353341,0.00065484346,0.00021632186,0.0020877148,0.0014891644,0.016996859,0.011405793,0.6476951,0.30058405,0.00010354234],"about_ca_topic_score_codex":0.002859385,"about_ca_topic_score_gemma":0.0033673984,"teacher_disagreement_score":0.014500676,"about_ca_system_score_codex":0.0022003266,"about_ca_system_score_gemma":0.003073863,"threshold_uncertainty_score":0.048509598},"labels":[],"label_agreement":null},{"id":"W2093127601","doi":"10.1002/spe.1013","title":"Spotting the difference","year":2010,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Readability; Usability; Interface (matter); Space (punctuation); Human–computer interaction; Software; User interface; Granularity; World Wide Web; Multimedia; Information retrieval; Programming language; Operating system","score_opus":0.014185157614998764,"score_gpt":0.2935626608430459,"score_spread":0.2793775032280471,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2093127601","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5864549,0.0023401328,0.3236363,0.007930331,0.004077278,0.0006901518,0.0011786756,0.019035108,0.05465706],"genre_scores_gemma":[0.7964247,0.00039183896,0.1858863,0.0016694149,0.00037344813,0.00022946509,0.00058704364,0.0014105416,0.013027368],"study_design_codex":"design_other","study_design_gemma":"design_other","domain_scores_codex":[0.9966742,0.0013847593,0.00025585882,0.0005203052,0.0010102042,0.000154827],"domain_scores_gemma":[0.97919047,0.013180901,0.0014738642,0.0027938264,0.0026662832,0.0006947457],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028165916,0.0008867342,0.0005832198,0.0014148585,0.00068984024,0.0021266635,0.0012189698,0.0016492617,0.02200908],"category_scores_gemma":[0.03293397,0.00024193533,0.0003336539,0.00074707833,0.0009369204,0.0044393674,0.0030793124,0.0012046592,0.004271393],"study_design_candidate":"design_other","study_design_consensus":"design_other","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026872011,0.00030999834,0.010110662,0.0012698531,0.000072658884,0.0019160901,0.015562015,0.00072661275,0.091151655,0.022510031,0.06699837,0.7866849],"study_design_scores_gemma":[0.00061746006,0.0038206093,0.05002577,0.0015381931,0.00038533573,0.013528293,0.021905314,0.052376177,0.19587804,0.07507806,0.5842727,0.0005740735],"about_ca_topic_score_codex":0.00020228619,"about_ca_topic_score_gemma":0.00021039402,"teacher_disagreement_score":0.02200908,"about_ca_system_score_codex":0.0003785484,"about_ca_system_score_gemma":0.00031952662,"threshold_uncertainty_score":0.07362771},"labels":[],"label_agreement":null},{"id":"W2093277130","doi":"10.1016/j.jss.2003.10.027","title":"Task-directed software inspection","year":2003,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software inspection; Suite; Task (project management); Software; Process (computing); Computer science; Visual inspection; Variety (cybernetics); Software engineering; Engineering; Risk analysis (engineering); Reliability engineering; Software development; Systems engineering; Software quality; Artificial intelligence; Business; Programming language","score_opus":0.013998482647934067,"score_gpt":0.2397085514835722,"score_spread":0.22571006883563813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2093277130","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16966756,0.00015848718,0.8130696,0.00020574237,0.000099417004,0.00017898185,0.00009830038,0.0054597547,0.011062232],"genre_scores_gemma":[0.84098595,0.00009366122,0.15014121,0.00011617043,0.000017873188,0.00007241762,0.0002149897,0.00024011663,0.008117637],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990339,0.0002937665,0.000043517455,0.00015797476,0.00034899177,0.00012187627],"domain_scores_gemma":[0.9954081,0.0017133745,0.000344689,0.0011336199,0.0011131589,0.00028695524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013162907,0.0005794731,0.00052885356,0.0006609827,0.00052930514,0.0005683613,0.0011700139,0.0008242335,0.0033161715],"category_scores_gemma":[0.006099071,0.00035868396,0.000390862,0.00036979123,0.00048374484,0.0006702841,0.0012696692,0.00078233646,0.0008198543],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015844008,0.0012559856,0.0113538345,0.00046572427,0.00010975147,0.0008643627,0.000950556,0.13367741,0.15674746,0.032286827,0.014599727,0.6461039],"study_design_scores_gemma":[0.00006966194,0.0004545918,0.00596634,0.00004460652,0.00006654286,0.00041808287,0.0001939329,0.9013972,0.048584837,0.035802983,0.0069521773,0.00004896382],"about_ca_topic_score_codex":0.0029110378,"about_ca_topic_score_gemma":0.005236678,"teacher_disagreement_score":0.0033161715,"about_ca_system_score_codex":0.0003318843,"about_ca_system_score_gemma":0.0012675432,"threshold_uncertainty_score":0.011093676},"labels":[],"label_agreement":null},{"id":"W2093400716","doi":"10.1007/s10664-015-9379-3","title":"What are mobile developers asking about? A large scale study using stack overflow","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":311,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Mobile device; World Wide Web; Mobile computing; Latent Dirichlet allocation; Context (archaeology); Software; Popularity; Mobile Web; Data science; Mobile technology; Software development; Topic model; Telecommunications; Artificial intelligence; Operating system","score_opus":0.0514213064798872,"score_gpt":0.32439770776851695,"score_spread":0.27297640128862977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2093400716","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99888664,0.000049347036,0.00016807002,0.00018334737,0.000004059008,0.00003792654,0.00005055216,0.000007002263,0.0006130934],"genre_scores_gemma":[0.9979048,0.00016268711,0.0005234713,0.00029065105,0.0000131345005,0.00012050925,0.000117254014,0.000018863422,0.0008487525],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9968951,0.001451067,0.00024268884,0.00035509077,0.0006447904,0.00041123878],"domain_scores_gemma":[0.90798396,0.06840722,0.010575872,0.002007544,0.006842412,0.00418298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057516918,0.0004887506,0.00043829912,0.0025253308,0.0031417918,0.0024425455,0.0008518639,0.0016634722,0.0018942838],"category_scores_gemma":[0.052636843,0.0007215972,0.0002750181,0.0021910467,0.0014851672,0.004181307,0.002157345,0.002351156,0.0005235501],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002817505,0.0021616563,0.7029215,0.00025769533,0.00006389122,0.0011832587,0.2563147,0.00011133351,0.0017727845,0.0007381536,0.0021138473,0.032079387],"study_design_scores_gemma":[0.00011568786,0.0009109539,0.6762496,0.00032482023,0.000116198695,0.0006376585,0.31242812,0.00087705394,0.0010933804,0.0005268846,0.006644589,0.00007508994],"about_ca_topic_score_codex":0.016766833,"about_ca_topic_score_gemma":0.040645234,"teacher_disagreement_score":0.016766833,"about_ca_system_score_codex":0.0018900449,"about_ca_system_score_gemma":0.0033947048,"threshold_uncertainty_score":0.033338487},"labels":[],"label_agreement":null},{"id":"W2093620756","doi":"10.1145/2593822.2593824","title":"Recommending a starting point for a programming task: an initial investigation","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Task (project management); Point (geometry); Programming language; Engineering; Mathematics; Systems engineering","score_opus":0.04741838711742813,"score_gpt":0.3210732864602837,"score_spread":0.2736548993428556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2093620756","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91854686,0.0040126843,0.06282627,0.001372105,0.0001318112,0.0011198121,0.0031928571,0.0012394481,0.007558188],"genre_scores_gemma":[0.8806171,0.001163412,0.109017365,0.00020093704,0.00004744179,0.00039545403,0.005303772,0.00021765254,0.0030368306],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99535745,0.0020998472,0.00036810365,0.0009910226,0.0009345083,0.00024904337],"domain_scores_gemma":[0.9050923,0.07718057,0.00213373,0.0036602595,0.010400606,0.0015325884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074672494,0.0008631019,0.0009893803,0.0030402278,0.0012185142,0.003302091,0.001411166,0.0017825587,0.0031806342],"category_scores_gemma":[0.06875881,0.0005511254,0.0009213076,0.0024325496,0.0006798445,0.005550735,0.0008327049,0.0021817603,0.0015953113],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044030813,0.004414461,0.43357807,0.0033637555,0.00070199807,0.0010271885,0.004086657,0.033367686,0.010490305,0.0038255635,0.02320833,0.47753292],"study_design_scores_gemma":[0.0007239254,0.005174515,0.29048327,0.001926215,0.0011714076,0.0013943205,0.012987952,0.61878717,0.01646503,0.008544125,0.041899003,0.00044315698],"about_ca_topic_score_codex":0.016520891,"about_ca_topic_score_gemma":0.03435996,"teacher_disagreement_score":0.016520891,"about_ca_system_score_codex":0.0015434523,"about_ca_system_score_gemma":0.002428078,"threshold_uncertainty_score":0.039491117},"labels":[],"label_agreement":null},{"id":"W2094303041","doi":"10.1109/re.2010.44","title":"Requirements Determination is Unstoppable: An Experience Report","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Focus (optics); Process (computing); Computer science; Software engineering; Requirements elicitation; Requirements engineering; Software; Engineering; Engineering management; Programming language","score_opus":0.03555206012098259,"score_gpt":0.3463870561952624,"score_spread":0.3108349960742798,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2094303041","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92476094,0.001913719,0.02869016,0.008950161,0.0002725202,0.00034500548,0.000086591754,0.0004890454,0.034491938],"genre_scores_gemma":[0.97263545,0.0018027343,0.010215951,0.0017382286,0.000082988976,0.00015552959,0.00010668747,0.0003491445,0.012913249],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9522789,0.031732015,0.0018317908,0.0018287397,0.010242077,0.0020865481],"domain_scores_gemma":[0.9119402,0.06575421,0.0032380035,0.0038765436,0.012084378,0.0031066579],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.037136085,0.0012955151,0.0008286605,0.0016141535,0.0070307655,0.004752874,0.0020705054,0.0029575387,0.0033016473],"category_scores_gemma":[0.106013305,0.0009705302,0.00057812646,0.0021496026,0.008173157,0.006839141,0.0053826272,0.004118833,0.0015706179],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015401156,0.00077304663,0.0047863815,0.00059901795,0.000025083498,0.0060828724,0.8696188,0.00039261073,0.0126624005,0.0048999786,0.0062491745,0.09375668],"study_design_scores_gemma":[0.000043672782,0.0015045314,0.007856972,0.0006081805,0.000060819148,0.015553634,0.7584863,0.0034391014,0.0155723635,0.003325425,0.19326587,0.00028312908],"about_ca_topic_score_codex":0.004266869,"about_ca_topic_score_gemma":0.006821554,"teacher_disagreement_score":0.037136085,"about_ca_system_score_codex":0.002860198,"about_ca_system_score_gemma":0.004033158,"threshold_uncertainty_score":0.19639677},"labels":[],"label_agreement":null},{"id":"W2094978367","doi":"10.1002/spip.273","title":"Release planning process improvement — an industrial case study","year":2006,"lang":"en","type":"article","venue":"Software Process Improvement and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Process (computing); Process management; Strategic planning; Engineering; Computer science; Business","score_opus":0.043721770110688916,"score_gpt":0.34822808937314986,"score_spread":0.30450631926246097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2094978367","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9771311,0.00018969276,0.01596534,0.0004728712,0.000018979616,0.0009122718,0.00008545052,0.00015519612,0.0050690933],"genre_scores_gemma":[0.963764,0.000229596,0.03394484,0.00007732945,0.000017523043,0.00041221082,0.00013640644,0.000029115652,0.001388881],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9932701,0.0042977,0.00027720275,0.00034343844,0.0013528162,0.00045873816],"domain_scores_gemma":[0.9717845,0.020310186,0.0021236634,0.0023997796,0.002605826,0.0007760204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00812535,0.00062653207,0.00047031033,0.0013843304,0.002058391,0.0014018046,0.002445944,0.0017531393,0.0018046362],"category_scores_gemma":[0.0158684,0.00043396436,0.0004605844,0.0019275905,0.0013892369,0.0009840543,0.0011566327,0.001408792,0.000399321],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0047446126,0.04213534,0.06192363,0.0023953333,0.00027077593,0.033801474,0.028624156,0.26659057,0.03079075,0.018787798,0.011751476,0.49818408],"study_design_scores_gemma":[0.006944748,0.04785372,0.07348792,0.00082522654,0.00051931257,0.014858286,0.031849068,0.6359121,0.10267248,0.010183766,0.074315645,0.0005777179],"about_ca_topic_score_codex":0.004459931,"about_ca_topic_score_gemma":0.00520446,"teacher_disagreement_score":0.00812535,"about_ca_system_score_codex":0.00218745,"about_ca_system_score_gemma":0.0020968684,"threshold_uncertainty_score":0.04297149},"labels":[],"label_agreement":null},{"id":"W2095485672","doi":"10.4236/jsea.2012.57060","title":"Empirical Analysis of Object-Oriented Design Metrics for Predicting Unit Testing Effort of Classes","year":2012,"lang":"en","type":"article","venue":"Journal of Software Engineering and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Innovation and Economic Development Trois Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Regression testing; Unit testing; Testability; Computer science; Logistic regression; Multivariate statistics; Metric (unit); Univariate; Non-regression testing; Cohesion (chemistry); Data mining; White-box testing; Keyword-driven testing; Reliability engineering; Machine learning; Software; Engineering; Programming language; Software development","score_opus":0.0563044086886836,"score_gpt":0.31747629891119916,"score_spread":0.2611718902225156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2095485672","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97666013,0.00035999296,0.021835575,0.00010046083,0.0000078230505,0.00003836777,0.0003340201,0.00011448395,0.000549195],"genre_scores_gemma":[0.9931416,0.000085267515,0.0060050944,0.0000137940515,0.000010020114,0.000051443072,0.0005638228,0.000028519722,0.00010034204],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99016273,0.005458621,0.0007496644,0.0007926345,0.0025111078,0.0003252593],"domain_scores_gemma":[0.6313781,0.31456584,0.034273926,0.00962658,0.008289435,0.001866155],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014213069,0.0010334344,0.00053273607,0.003941159,0.00019346223,0.00082045415,0.00085339166,0.00082361995,0.0005627198],"category_scores_gemma":[0.1607009,0.00026575328,0.00061219773,0.0025843643,0.0004820197,0.0017084796,0.0005968698,0.00088335725,0.00030076542],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001319143,0.00020913706,0.9613918,0.000056600522,0.0001843308,0.00006620958,0.0001784894,0.015028093,0.0006897552,0.00027095774,0.00017784254,0.021614814],"study_design_scores_gemma":[0.000031897045,0.0005337805,0.79430944,0.00005355348,0.000096758966,0.00027978118,0.0001593724,0.2007715,0.002225624,0.0009972542,0.0005081303,0.00003306888],"about_ca_topic_score_codex":0.00144791,"about_ca_topic_score_gemma":0.0015356496,"teacher_disagreement_score":0.9857869,"about_ca_system_score_codex":0.00046706753,"about_ca_system_score_gemma":0.000520786,"threshold_uncertainty_score":0.07516682},"labels":[],"label_agreement":null},{"id":"W2095719978","doi":"10.1109/csmr.2008.4493323","title":"Use Case Redocumentation from GUI Event Traces","year":2008,"lang":"en","type":"article","venue":"Proceedings of the ... European Conference on Software Maintenance and Reengineering/Proceedings of the European Conference on Software Maintenance and Reengineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Ministry of Advanced Education, Government of Alberta","keywords":"Computer science; Documentation; Task (project management); Java; Event (particle physics); Software architecture; Software engineering; Interface (matter); Graphical user interface; Software; User interface; Programming language; Operating system; Engineering","score_opus":0.029524690256025724,"score_gpt":0.22896880737659706,"score_spread":0.19944411712057133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2095719978","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23056847,0.0008082917,0.70977354,0.0004384866,0.00019065627,0.0015081799,0.011334227,0.03708322,0.008294966],"genre_scores_gemma":[0.38911936,0.0006946183,0.57847595,0.00006290672,0.0000785544,0.00094957923,0.022106707,0.0028767555,0.0056356243],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962812,0.0005490128,0.0004764777,0.0008377554,0.0016661257,0.00018935431],"domain_scores_gemma":[0.970214,0.010235414,0.004098557,0.00783761,0.006786067,0.0008283932],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027664367,0.0016152959,0.0010110924,0.012190273,0.00083232013,0.0031829814,0.001895385,0.0011197388,0.0029905206],"category_scores_gemma":[0.033257127,0.00074770267,0.00074680324,0.0054806126,0.0004955291,0.0019712753,0.0025420878,0.0014857535,0.0014161809],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006513352,0.0006099569,0.046824746,0.0010640044,0.00023880858,0.0022529345,0.008516323,0.0144836875,0.043780345,0.007112577,0.015225457,0.85923976],"study_design_scores_gemma":[0.00019373259,0.00065593666,0.14051111,0.0010440107,0.00054292474,0.0046011535,0.007049903,0.49072728,0.18938555,0.030036164,0.13459785,0.00065446045],"about_ca_topic_score_codex":0.0066817147,"about_ca_topic_score_gemma":0.009114554,"teacher_disagreement_score":0.012190273,"about_ca_system_score_codex":0.0008192674,"about_ca_system_score_gemma":0.0016472458,"threshold_uncertainty_score":0.014630437},"labels":[],"label_agreement":null},{"id":"W2095743800","doi":"10.1109/icdm.2005.120","title":"Predicting Software Escalations with Maximum ROI","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Vendor; Return on investment; Reputation; Revenue; Software; Computer science; Investment (military); Product (mathematics); Business; Marketing; Finance; Operating system","score_opus":0.008598635395594132,"score_gpt":0.2201097054914591,"score_spread":0.21151107009586498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2095743800","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8695662,0.0005219253,0.124861,0.0005987048,0.000032309206,0.00009244555,0.0014627109,0.0012610474,0.0016034881],"genre_scores_gemma":[0.97348255,0.0001047839,0.024245523,0.000040918723,0.000029942144,0.00004097863,0.0016403454,0.000029619738,0.00038539886],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983437,0.0004780792,0.00015145274,0.0003856975,0.00048164552,0.00015933838],"domain_scores_gemma":[0.98744255,0.007107257,0.0026437137,0.0010000815,0.0014373335,0.00036904658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032014726,0.00076379615,0.00071176194,0.0033124445,0.00024378498,0.000868134,0.0009440099,0.0011222686,0.00055252394],"category_scores_gemma":[0.017176338,0.00034932303,0.0005325023,0.0015102194,0.0002842544,0.0016596058,0.00056340086,0.0010559246,0.00031614353],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048609407,0.00050016685,0.259965,0.0001335306,0.00019311241,0.0003088632,0.00015667331,0.53525585,0.0026012873,0.001737784,0.005456394,0.1932052],"study_design_scores_gemma":[0.000008471239,0.00007903306,0.015372536,0.0000078970415,0.000019238232,0.00009282897,0.00003010926,0.9811757,0.0013179195,0.0015974987,0.0002894399,0.0000092423015],"about_ca_topic_score_codex":0.0020696563,"about_ca_topic_score_gemma":0.0023387792,"teacher_disagreement_score":0.0033124445,"about_ca_system_score_codex":0.0007523882,"about_ca_system_score_gemma":0.00038294334,"threshold_uncertainty_score":0.016931176},"labels":[],"label_agreement":null},{"id":"W2095802649","doi":"10.1145/1251535.1251542","title":"Comparing call graphs","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Call graph; Computer science; Spurious relationship; Graph; Theoretical computer science; Power graph analysis; Data mining; Machine learning","score_opus":0.03137187158833518,"score_gpt":0.2839632137052097,"score_spread":0.2525913421168745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2095802649","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25913313,0.0016380433,0.68017995,0.0010989183,0.0002309824,0.00066901796,0.008071318,0.019403102,0.0295755],"genre_scores_gemma":[0.7139592,0.00081910787,0.26507202,0.00032287065,0.00008559905,0.0004934078,0.01088355,0.004405617,0.00395867],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99221945,0.0021053988,0.00061513943,0.00096568186,0.0036937508,0.0004005716],"domain_scores_gemma":[0.95837146,0.028502705,0.0023122248,0.0043691453,0.0060749426,0.00036954015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003945396,0.0010024828,0.00072801183,0.009844533,0.0011005381,0.0032312432,0.0015811507,0.0011419147,0.0066139298],"category_scores_gemma":[0.04092423,0.0004210261,0.0013560634,0.0061310623,0.0009856387,0.004180326,0.0017469933,0.0010537013,0.0010094172],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013836374,0.00028653,0.052086342,0.0028200669,0.0007727826,0.0010766576,0.003401403,0.076868765,0.037514452,0.08266948,0.020896371,0.7202235],"study_design_scores_gemma":[0.00023416031,0.0009251271,0.09233014,0.00068897154,0.0013446861,0.0023492058,0.0054607596,0.3476343,0.078831255,0.33559835,0.13404159,0.0005614549],"about_ca_topic_score_codex":0.004654593,"about_ca_topic_score_gemma":0.004171021,"teacher_disagreement_score":0.009844533,"about_ca_system_score_codex":0.0015318649,"about_ca_system_score_gemma":0.001480719,"threshold_uncertainty_score":0.02212584},"labels":[],"label_agreement":null},{"id":"W2095867894","doi":"10.11575/prism/31259","title":"Frequently Asked Questions in Bug Reports","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Eclipse; Software bug; Data science; Software; World Wide Web; Programming language","score_opus":0.012796010793236388,"score_gpt":0.2764901512939575,"score_spread":0.2636941405007211,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2095867894","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9711737,0.0017042164,0.01794914,0.0011770194,0.00011271828,0.00071854546,0.0020491,0.00028326377,0.0048322617],"genre_scores_gemma":[0.9752734,0.0005748887,0.018468164,0.0005045665,0.00006274143,0.0016019995,0.0019912075,0.000104748746,0.0014182477],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9203476,0.057886064,0.007794598,0.0035338355,0.008077906,0.0023599786],"domain_scores_gemma":[0.67853457,0.26017573,0.029732514,0.0073720277,0.021064503,0.0031206587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028669104,0.00077262544,0.0010518663,0.006780603,0.0015922883,0.002559915,0.0011145248,0.0028258692,0.0032378647],"category_scores_gemma":[0.1820015,0.0005390738,0.0007537952,0.005399348,0.0018666339,0.0038832277,0.004375092,0.0014224749,0.00089064165],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014932649,0.00051197514,0.3977463,0.00406835,0.0003237789,0.0013729628,0.3929778,0.002344866,0.008673313,0.0093811015,0.009624415,0.17148186],"study_design_scores_gemma":[0.00019350031,0.0017102794,0.5903538,0.0025140976,0.0002589545,0.003045347,0.2681364,0.008026062,0.007609749,0.019448172,0.09816973,0.00053386553],"about_ca_topic_score_codex":0.0017760424,"about_ca_topic_score_gemma":0.0013249062,"teacher_disagreement_score":0.028669104,"about_ca_system_score_codex":0.002488398,"about_ca_system_score_gemma":0.0009650013,"threshold_uncertainty_score":0.15161854},"labels":[],"label_agreement":null},{"id":"W2096211673","doi":"10.1109/wcre.2012.23","title":"SCAN: An Approach to Label and Relate Execution Trace Segments","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; TRACE (psycholinguistics); Program comprehension; Task (project management); Artificial intelligence; Formal concept analysis; Natural language processing; Data mining; Information retrieval; Programming language; Software; Algorithm","score_opus":0.03396747430075324,"score_gpt":0.28240211769735224,"score_spread":0.248434643396599,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096211673","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070261615,0.00013889666,0.9734838,0.00021492662,0.000052931224,0.0005911968,0.002601251,0.01346618,0.0024246997],"genre_scores_gemma":[0.026790693,0.00011285375,0.96558785,0.00009087767,0.000026242174,0.0005256424,0.0033880952,0.00091646134,0.002561424],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99527884,0.00094224303,0.00048120937,0.0012950826,0.0017753771,0.0002272062],"domain_scores_gemma":[0.9873087,0.004913167,0.0018850365,0.0023616299,0.003067855,0.00046358848],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003810012,0.0017024209,0.00104211,0.0114765465,0.0015404469,0.0032307366,0.003036117,0.001884546,0.0085683465],"category_scores_gemma":[0.016635332,0.00081164704,0.0016573683,0.0063458243,0.001572695,0.005531725,0.0037336745,0.0024060304,0.0034030427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009846936,0.00048150108,0.013208074,0.0017612444,0.00028641478,0.0008812989,0.008319217,0.012417408,0.041036334,0.047721285,0.030547997,0.84235454],"study_design_scores_gemma":[0.00019735972,0.0009598779,0.011512267,0.0008323518,0.00039628154,0.0021423993,0.006937575,0.46428183,0.08490061,0.117726915,0.30972028,0.00039219361],"about_ca_topic_score_codex":0.007861868,"about_ca_topic_score_gemma":0.01301723,"teacher_disagreement_score":0.0114765465,"about_ca_system_score_codex":0.001284915,"about_ca_system_score_gemma":0.0030626508,"threshold_uncertainty_score":0.028663993},"labels":[],"label_agreement":null},{"id":"W2096370414","doi":"10.1109/ase.2009.89","title":"Evaluating the Accuracy of Fault Localization Techniques","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Data mining; Process (computing); Assertion; Artificial intelligence; Fault (geology); Class (philosophy); Machine learning; Programming language","score_opus":0.05549973097779415,"score_gpt":0.3941240675486642,"score_spread":0.33862433657087004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096370414","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61713856,0.0049778204,0.358328,0.002244965,0.0002789664,0.00026783635,0.0016315622,0.0065679518,0.008564296],"genre_scores_gemma":[0.884936,0.0006028304,0.11162465,0.00023020481,0.00008712527,0.000099765086,0.0013930487,0.00029191506,0.0007345298],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9425495,0.018423373,0.006431099,0.0056795287,0.0250791,0.0018373597],"domain_scores_gemma":[0.6666535,0.24472304,0.02219452,0.03493897,0.030218305,0.0012716533],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031839345,0.0015644603,0.0017450966,0.009959233,0.0007651304,0.0029908859,0.0029132776,0.0030539017,0.0009770753],"category_scores_gemma":[0.20071049,0.0005180684,0.0017458193,0.0052759433,0.0014700553,0.0053044655,0.0017261187,0.0017104392,0.00094248395],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017528174,0.000684806,0.23048049,0.0013329092,0.0015885648,0.00031575834,0.0008746974,0.16639562,0.015187178,0.006768191,0.0062953504,0.5683236],"study_design_scores_gemma":[0.00017612743,0.0020570706,0.06439114,0.0003096603,0.00063381996,0.0007056966,0.0006103001,0.85248584,0.06167388,0.011209273,0.0055881147,0.00015913138],"about_ca_topic_score_codex":0.0029894887,"about_ca_topic_score_gemma":0.0028822827,"teacher_disagreement_score":0.031839345,"about_ca_system_score_codex":0.0014527258,"about_ca_system_score_gemma":0.0011598741,"threshold_uncertainty_score":0.16838455},"labels":[],"label_agreement":null},{"id":"W2096450604","doi":"10.1109/wpc.2004.1311061","title":"An effectiveness measure for software clustering algorithms","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":152,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Computer science; Measure (data warehouse); Data mining; Software; Algorithm; Software metric; Machine learning; Software system; Artificial intelligence; Software construction","score_opus":0.02372367782644497,"score_gpt":0.2955912787144825,"score_spread":0.2718676008880375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096450604","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05452799,0.005733771,0.92820907,0.00068692834,0.0002857888,0.00042962367,0.00042671393,0.0011115909,0.0085885],"genre_scores_gemma":[0.39709345,0.0015109086,0.59735894,0.00020320601,0.0006641386,0.0005662477,0.0009933925,0.00036560104,0.0012440669],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97669876,0.007542032,0.0020058958,0.0015628637,0.011650648,0.000539819],"domain_scores_gemma":[0.91539234,0.061685264,0.004620328,0.004268925,0.012581147,0.001452048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013515006,0.0016770733,0.0023471855,0.014462077,0.0017612275,0.003486092,0.0018095977,0.0023561777,0.0013786461],"category_scores_gemma":[0.073735505,0.0004630889,0.0014009399,0.005310691,0.002191539,0.00640722,0.0015986802,0.0016915551,0.00065879227],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007990787,0.0005968747,0.029691998,0.0018500598,0.0009892996,0.00023658822,0.0006743641,0.17868823,0.02023424,0.09910925,0.01076559,0.65636444],"study_design_scores_gemma":[0.00011261934,0.0018167453,0.01725279,0.0002861012,0.0004476188,0.0012076204,0.0005577878,0.8630191,0.030468402,0.06632325,0.018305399,0.00020250362],"about_ca_topic_score_codex":0.0008677412,"about_ca_topic_score_gemma":0.0007965401,"teacher_disagreement_score":0.014462077,"about_ca_system_score_codex":0.002497679,"about_ca_system_score_gemma":0.0012292718,"threshold_uncertainty_score":0.07147509},"labels":[],"label_agreement":null},{"id":"W2096591909","doi":"10.1002/smr.311","title":"Software Maintenance Maturity Model (SMmm): the software maintenance process model","year":2005,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":146,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Capability Maturity Model Integration; Capability Maturity Model; Software maintenance; LeanCMMI; Software engineering; Scope (computer science); Computer science; Software development; Process management; Systems engineering; Software development process; Engineering; Software; Operating system","score_opus":0.04068226693628353,"score_gpt":0.33987777925729934,"score_spread":0.2991955123210158,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096591909","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01799098,0.006997102,0.9197076,0.012454879,0.0001943253,0.0010749111,0.00083267636,0.0016933546,0.039054133],"genre_scores_gemma":[0.28894228,0.005318986,0.69559014,0.0012421034,0.00017381733,0.0023604152,0.001569671,0.00014960502,0.004652929],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9917904,0.0029178385,0.0009531793,0.0007142932,0.0031660155,0.0004583027],"domain_scores_gemma":[0.98662335,0.005536968,0.0021867733,0.0009232337,0.0040822523,0.00064739166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011256055,0.0011490044,0.0007529158,0.004069388,0.0013175186,0.0050822403,0.0027302552,0.004206402,0.0015353331],"category_scores_gemma":[0.02793576,0.0005557903,0.0013563095,0.0039281654,0.0017488577,0.009097503,0.0026783694,0.0034103943,0.0011926015],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000060173268,0.00020362898,0.008475148,0.0013524139,0.00015058827,0.00028090322,0.00399599,0.03602974,0.0031616956,0.6526524,0.014230244,0.2794072],"study_design_scores_gemma":[0.0000947924,0.00045993383,0.007359082,0.0023203508,0.00027944913,0.0009958921,0.001155988,0.12102578,0.0034400905,0.6525304,0.21013825,0.00019997038],"about_ca_topic_score_codex":0.0056721694,"about_ca_topic_score_gemma":0.004036701,"teacher_disagreement_score":0.011256055,"about_ca_system_score_codex":0.00489553,"about_ca_system_score_gemma":0.010334593,"threshold_uncertainty_score":0.05952841},"labels":[],"label_agreement":null},{"id":"W2096687600","doi":"10.1109/isese.2004.7","title":"An empirical study of a Qualitative Systematic Approach to Requirements Analysis (QSARA)","year":2004,"lang":"en","type":"article","venue":"UTS ePRESS (University of Technology Sydney)","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Construct (python library); Empirical research; Interdependence; Domain (mathematical analysis); Management science; Systematic review; Requirements analysis; Knowledge management; Data science; Engineering","score_opus":0.06371745272885647,"score_gpt":0.35091770736320926,"score_spread":0.2872002546343528,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096687600","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9041466,0.00024317384,0.06000683,0.0018340081,0.00011301896,0.022537295,0.00078256906,0.00008962141,0.010246748],"genre_scores_gemma":[0.8346476,0.00055128406,0.12367837,0.001538965,0.00006523641,0.033617403,0.0004891448,0.000090064976,0.0053219735],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9421573,0.046408903,0.0019441706,0.0023048008,0.004533699,0.0026511392],"domain_scores_gemma":[0.8121827,0.14893605,0.0056258915,0.005698185,0.022970255,0.0045868303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.055230454,0.0008376828,0.0010283848,0.003031038,0.007144142,0.002504298,0.0015193904,0.0018079146,0.005191885],"category_scores_gemma":[0.11178434,0.0011459214,0.0005507966,0.0032086326,0.004683745,0.0029648354,0.0030885963,0.0022299339,0.0010683236],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013579628,0.008289889,0.024620766,0.0033756222,0.000042508684,0.0028761881,0.75201327,0.002796503,0.017598666,0.012566807,0.0042119757,0.1702499],"study_design_scores_gemma":[0.0007985105,0.00996393,0.018327681,0.0019943116,0.00006596551,0.0009988962,0.88059735,0.006632137,0.011989097,0.0064770803,0.061958034,0.0001969911],"about_ca_topic_score_codex":0.0037320482,"about_ca_topic_score_gemma":0.009531867,"teacher_disagreement_score":0.055230454,"about_ca_system_score_codex":0.0065850476,"about_ca_system_score_gemma":0.013599608,"threshold_uncertainty_score":0.29209006},"labels":[],"label_agreement":null},{"id":"W2096783995","doi":"10.5555/2337223.2337272","title":"Automated analysis of CSS rules to support style maintenance","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Cascading Style Sheets; Web application; Inheritance (genetic algorithm); Semantic Web; Semantics (computer science); Class (philosophy); World Wide Web; Software engineering; Programming language; Web page; Artificial intelligence","score_opus":0.01899203707600168,"score_gpt":0.29630894706467764,"score_spread":0.277316909988676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096783995","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23105484,0.00047433382,0.69108343,0.00025233158,0.00008632884,0.00038993126,0.0019372329,0.07117954,0.0035420195],"genre_scores_gemma":[0.5181544,0.00018009717,0.47515282,0.00007070209,0.000043704764,0.00013844724,0.0030940142,0.0015376759,0.0016280117],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974082,0.00056094385,0.00030754987,0.00056508655,0.0010383391,0.000119821445],"domain_scores_gemma":[0.97424453,0.010673818,0.0031516864,0.004929293,0.006648108,0.00035258709],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018962078,0.0009745553,0.00063818216,0.0063100904,0.00066558464,0.0017488685,0.0013822372,0.0006422156,0.0018664019],"category_scores_gemma":[0.016342413,0.00040380438,0.0006063895,0.0021755986,0.00041438188,0.0010293395,0.0008345457,0.0007474056,0.0014794095],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041065214,0.00043429987,0.075049974,0.0004684004,0.0001499549,0.0006685666,0.0010020948,0.016281234,0.084282696,0.0021440436,0.012333761,0.8067743],"study_design_scores_gemma":[0.00006564235,0.00015440119,0.02771837,0.00012517259,0.00012387261,0.0011057573,0.00022793355,0.78541,0.16670866,0.0044365167,0.01381781,0.00010593405],"about_ca_topic_score_codex":0.0033245604,"about_ca_topic_score_gemma":0.0038114672,"teacher_disagreement_score":0.0063100904,"about_ca_system_score_codex":0.00061933027,"about_ca_system_score_gemma":0.0013135824,"threshold_uncertainty_score":0.010028183},"labels":[],"label_agreement":null},{"id":"W2096807669","doi":"10.1109/wcre.2008.25","title":"A Business Process Explorer: Recovering Business Processes from Business Applications","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Artifact-centric business process model; Business process modeling; Business process discovery; Business rule; Business process; Business Process Model and Notation; Business process management; Computer science; Business analysis; Process management; Business domain; Business requirements; Business activity monitoring; New business development; Process (computing); Business object; Business model; Business; Work in process; Marketing; Operating system","score_opus":0.03310749556323077,"score_gpt":0.2583279268951812,"score_spread":0.2252204313319504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096807669","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022024075,0.00046758703,0.8996417,0.00042172588,0.000037782356,0.00029337776,0.0012933093,0.074523665,0.0012969337],"genre_scores_gemma":[0.07567324,0.0004193951,0.91112375,0.00019142544,0.000036263013,0.00024918202,0.004434622,0.0048706904,0.0030014224],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972217,0.0005542437,0.00020839175,0.0005678874,0.0012979042,0.00014983423],"domain_scores_gemma":[0.9891216,0.0068116505,0.0012309692,0.0018626953,0.0006990345,0.00027400968],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003015385,0.0017960202,0.0010263489,0.0039891093,0.00086206145,0.0023071778,0.0022714345,0.0025766005,0.003763196],"category_scores_gemma":[0.017047258,0.0011013892,0.0018773878,0.0017457079,0.0011239724,0.0046239975,0.003389466,0.002622678,0.0026094923],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075968454,0.00041750187,0.00852736,0.0018634269,0.00023738356,0.004418649,0.0039425488,0.014440484,0.08806202,0.016188376,0.02825527,0.83288723],"study_design_scores_gemma":[0.0003554724,0.00055328134,0.009894192,0.0005781068,0.0002536982,0.008216515,0.0013409863,0.5268446,0.22918314,0.048522178,0.17389162,0.00036624004],"about_ca_topic_score_codex":0.0013625083,"about_ca_topic_score_gemma":0.0014641401,"teacher_disagreement_score":0.0039891093,"about_ca_system_score_codex":0.0004327907,"about_ca_system_score_gemma":0.0013851917,"threshold_uncertainty_score":0.015947044},"labels":[],"label_agreement":null},{"id":"W2096881356","doi":"10.1109/fuzz.2003.1209439","title":"Fuzzy clustering of software metrics","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Cluster analysis; Unsupervised learning; Software metric; Data mining; Software quality; Software; Machine learning; Software construction; Artificial intelligence; Domain (mathematical analysis); Software sizing; Software measurement; Software development; Programming language; Mathematics","score_opus":0.02371649863143328,"score_gpt":0.26711149989314575,"score_spread":0.24339500126171248,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096881356","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14224361,0.00041446343,0.85363555,0.00021088668,0.000027036307,0.0001052807,0.00015454982,0.00023006056,0.0029785172],"genre_scores_gemma":[0.7098036,0.00025014265,0.28761137,0.00006490926,0.00004465707,0.000100299956,0.00053625245,0.00007692956,0.0015117214],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971944,0.0009469623,0.00013075097,0.00056207576,0.000996058,0.00016976989],"domain_scores_gemma":[0.9878419,0.0068091922,0.0013385556,0.0011719203,0.0026063765,0.00023203205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031064923,0.00073927204,0.0008748598,0.0047053345,0.0009814467,0.0015432289,0.0013106526,0.00095969613,0.0007922367],"category_scores_gemma":[0.018804308,0.00032191188,0.00097295165,0.003368,0.0012530616,0.001755541,0.000915342,0.00082018133,0.0002711559],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017507712,0.0001721115,0.014210984,0.00021771087,0.00024324094,0.00015187888,0.0009984479,0.67544687,0.010164272,0.077868715,0.001785509,0.21856517],"study_design_scores_gemma":[0.00000701838,0.00005716641,0.00450665,0.000014510765,0.0000152259145,0.00006339154,0.00010599899,0.9515129,0.0029926917,0.03971234,0.000978803,0.00003325332],"about_ca_topic_score_codex":0.008454837,"about_ca_topic_score_gemma":0.0060249665,"teacher_disagreement_score":0.008454837,"about_ca_system_score_codex":0.0019144971,"about_ca_system_score_gemma":0.0009999722,"threshold_uncertainty_score":0.016811252},"labels":[],"label_agreement":null},{"id":"W2096960709","doi":"10.1109/icsm.2015.7332503","title":"Crowdsourced bug triaging","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Exploit; Computer science; Context (archaeology); Software bug; Task (project management); Data science; Software engineering; Computer security; Software; Programming language; Systems engineering; Engineering","score_opus":0.05216798005183396,"score_gpt":0.28988105184109414,"score_spread":0.23771307178926018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096960709","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20470682,0.00598795,0.69521755,0.0032980181,0.0025761542,0.0039925417,0.017140687,0.027630419,0.039449923],"genre_scores_gemma":[0.628896,0.00097203976,0.33260715,0.0010119401,0.00072700006,0.0018147334,0.015506243,0.001273711,0.017191201],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9907334,0.0031132437,0.00041471692,0.0028301484,0.002498026,0.00041051826],"domain_scores_gemma":[0.97384745,0.013530889,0.002181332,0.0058096647,0.0036875685,0.0009431091],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007352547,0.0021039885,0.002173766,0.008316281,0.001896531,0.0018382209,0.0031571155,0.00219469,0.00593889],"category_scores_gemma":[0.03281031,0.0007492454,0.0013471928,0.0061068544,0.00090794015,0.0019927362,0.0054004854,0.0014323744,0.004105205],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000714014,0.00075346004,0.035966024,0.0019301891,0.00076643954,0.001025965,0.0025735227,0.043916725,0.009243183,0.004858941,0.10844071,0.7898108],"study_design_scores_gemma":[0.000803745,0.00083587534,0.060506783,0.00086136744,0.00073478254,0.0015941014,0.0035457616,0.6042227,0.015843675,0.07779997,0.23277554,0.00047571427],"about_ca_topic_score_codex":0.011909421,"about_ca_topic_score_gemma":0.02123018,"teacher_disagreement_score":0.011909421,"about_ca_system_score_codex":0.0009752891,"about_ca_system_score_gemma":0.00263436,"threshold_uncertainty_score":0.03888446},"labels":[],"label_agreement":null},{"id":"W2097056679","doi":"10.1145/1985793.1986031","title":"Build system maintenance","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software maintenance; Computer science; Codebase; Software engineering; Overhead (engineering); Source code; Software development; Legacy system; Deliverable; Restructuring; Software; Systems engineering; Operating system; Engineering; Business","score_opus":0.025520968124238277,"score_gpt":0.2289150315217008,"score_spread":0.20339406339746252,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097056679","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3104211,0.002383635,0.3982972,0.0034812412,0.0009580561,0.004669479,0.009563744,0.06940978,0.20081569],"genre_scores_gemma":[0.65939546,0.0014434012,0.21640396,0.0015317991,0.00026288696,0.0017056445,0.018401813,0.008296918,0.092558146],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99158746,0.0011394455,0.00067552796,0.0010950621,0.0048545334,0.0006479467],"domain_scores_gemma":[0.95215,0.009768697,0.0049294666,0.019980308,0.011625402,0.00154615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067904107,0.0008664052,0.0005280778,0.0029047027,0.0017072489,0.0028625545,0.0030066404,0.0011185692,0.021805584],"category_scores_gemma":[0.04554671,0.0008136886,0.00065674016,0.001688875,0.000980026,0.004452339,0.0039060218,0.0016605804,0.012073766],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045515352,0.00077380595,0.05112754,0.0015528855,0.00013770352,0.0010592159,0.005576347,0.006980186,0.019150887,0.024268843,0.12053725,0.7683802],"study_design_scores_gemma":[0.00020272772,0.0008289369,0.069962494,0.0005387497,0.00021712907,0.0040137237,0.0016346844,0.02337199,0.029654,0.012841491,0.8565527,0.00018135742],"about_ca_topic_score_codex":0.0034012285,"about_ca_topic_score_gemma":0.0037339209,"teacher_disagreement_score":0.021805584,"about_ca_system_score_codex":0.0015362472,"about_ca_system_score_gemma":0.00357513,"threshold_uncertainty_score":0.072946966},"labels":[],"label_agreement":null},{"id":"W2097186696","doi":"10.1109/icst.2008.57","title":"An Empirical Study on Bayesian Network-based Approach for Test Case Prioritization","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Regression testing; Bayesian network; Java; Test case; Data mining; Software bug; Machine learning; Bayesian probability; Prioritization; Empirical research; Test (biology); Software; Artificial intelligence; Regression analysis; Software system; Engineering; Software construction","score_opus":0.05057616437589533,"score_gpt":0.3312695450053025,"score_spread":0.28069338062940713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097186696","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8182941,0.001441489,0.16890797,0.0012507796,0.000051617128,0.0005112494,0.0006199679,0.00042578529,0.008497023],"genre_scores_gemma":[0.93466616,0.00033925998,0.06353194,0.000099204655,0.000027839526,0.00018673835,0.00052085734,0.000049301183,0.0005787376],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96734256,0.02271826,0.000976223,0.0015876931,0.0069378675,0.00043731194],"domain_scores_gemma":[0.47395608,0.4902244,0.012021463,0.007995819,0.014521619,0.001280637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03785658,0.0007914934,0.00048295697,0.0033565036,0.0005186979,0.0011699453,0.0017254798,0.0013436895,0.0024252164],"category_scores_gemma":[0.26806694,0.00041707951,0.00035486,0.0032420333,0.000783032,0.0038298466,0.0008230124,0.0021392533,0.0002652882],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032937564,0.0038451783,0.2820927,0.0010278638,0.00055185956,0.00031258768,0.0016954822,0.22342896,0.0044784406,0.015379302,0.004222806,0.45967102],"study_design_scores_gemma":[0.0001899126,0.001097105,0.07610572,0.00018069349,0.00015671623,0.00044385655,0.0006447745,0.9084056,0.002813474,0.0067307744,0.003155517,0.00007579971],"about_ca_topic_score_codex":0.0070432466,"about_ca_topic_score_gemma":0.0080795325,"teacher_disagreement_score":0.03785658,"about_ca_system_score_codex":0.002477035,"about_ca_system_score_gemma":0.0012958909,"threshold_uncertainty_score":0.20020711},"labels":[],"label_agreement":null},{"id":"W2097510100","doi":"10.1109/wcre.2010.35","title":"From Whence It Came: Detecting Source Code Clones by Analyzing Assembler","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; CA Technologies","keywords":"Computer science; Source code; Programming language; Java; Normalization (sociology); Security token; Code (set theory); Open source; Operating system; Software","score_opus":0.01478714588735309,"score_gpt":0.2727150354436737,"score_spread":0.2579278895563206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097510100","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46763477,0.00039681586,0.50960314,0.00038076233,0.000079858335,0.00007202898,0.0004370406,0.017391933,0.004003662],"genre_scores_gemma":[0.74054825,0.00020816855,0.25289148,0.00026109986,0.000035445242,0.00004828006,0.0008515999,0.0024105788,0.002744956],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983998,0.00023380897,0.000091642156,0.0003206178,0.0008333258,0.0001207627],"domain_scores_gemma":[0.9939898,0.001722421,0.0010573696,0.0013276478,0.0017359989,0.00016675444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010224646,0.0005241193,0.00047197743,0.0026478765,0.0005895647,0.0014321036,0.00065756997,0.0006977032,0.0007742063],"category_scores_gemma":[0.008995345,0.0003438716,0.0004672129,0.0011647147,0.0008673283,0.0021303208,0.0010664089,0.0008995857,0.0007446816],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006070129,0.00020914759,0.118978195,0.00028789396,0.00014020156,0.0030470612,0.0047602165,0.009477037,0.30455828,0.017919071,0.0041193035,0.5358966],"study_design_scores_gemma":[0.000034434244,0.00042891048,0.064477146,0.000118996846,0.00020691741,0.004023698,0.0010657154,0.21313629,0.6683579,0.019252533,0.028625263,0.00027222472],"about_ca_topic_score_codex":0.001432349,"about_ca_topic_score_gemma":0.0016346563,"teacher_disagreement_score":0.0026478765,"about_ca_system_score_codex":0.00052388426,"about_ca_system_score_gemma":0.00070807984,"threshold_uncertainty_score":0.0054073334},"labels":[],"label_agreement":null},{"id":"W2097608484","doi":"10.1109/icse-companion.2009.5071005","title":"Mining recurrent activities: Fourier analysis of change events","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Fourier analysis; Fourier transform; Computer science; Software; Fourier series; Fourier sine and cosine series; Sine wave; SIGNAL (programming language); Signal processing; Field (mathematics); Data mining; Time series; Process (computing); Algorithm; Mathematics; Physics; Digital signal processing; Mathematical analysis; Machine learning; Fractional Fourier transform","score_opus":0.0706121917936518,"score_gpt":0.3246916657847333,"score_spread":0.2540794739910815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097608484","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40052837,0.0013784608,0.58606386,0.00073891674,0.00015132256,0.00021823224,0.002716868,0.0022197235,0.0059841815],"genre_scores_gemma":[0.884236,0.00094261294,0.10877605,0.000072548064,0.0001859217,0.00015991495,0.003704229,0.00013087015,0.0017918348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99881303,0.00018893101,0.00012476841,0.00029269466,0.0004525538,0.0001279979],"domain_scores_gemma":[0.9968015,0.0014742365,0.00065533025,0.00039070103,0.0005764113,0.00010172819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012832737,0.0006168822,0.00070718187,0.006396771,0.00035368177,0.0011266876,0.0007935847,0.0007477653,0.0015743191],"category_scores_gemma":[0.0076198583,0.00026184958,0.0007871497,0.004503777,0.00040390718,0.0014427623,0.0007448246,0.00066144613,0.00083063537],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005736592,0.0004723114,0.08931761,0.0005243294,0.00031226838,0.0019189329,0.0014011967,0.037171967,0.03234956,0.0124934465,0.008123658,0.8153409],"study_design_scores_gemma":[0.00006754165,0.0002947368,0.13273484,0.00015173326,0.00024088888,0.0024488976,0.001215383,0.7913187,0.0233472,0.03183287,0.016215997,0.00013127839],"about_ca_topic_score_codex":0.002424702,"about_ca_topic_score_gemma":0.0017172447,"teacher_disagreement_score":0.006396771,"about_ca_system_score_codex":0.00037347132,"about_ca_system_score_gemma":0.00039210383,"threshold_uncertainty_score":0.006786704},"labels":[],"label_agreement":null},{"id":"W2097611699","doi":"10.1109/wcre.2003.1287258","title":"An industrial experience in reverse engineering","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Reverse engineering; Computer science; Manufacturing engineering; Engineering; Programming language","score_opus":0.04825557935711727,"score_gpt":0.2937931619264098,"score_spread":0.24553758256929253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097611699","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5425169,0.0096076,0.24982485,0.025300575,0.0009222569,0.0003623471,0.00019998687,0.001511584,0.16975392],"genre_scores_gemma":[0.83342844,0.0061905608,0.11530905,0.005609192,0.00030315874,0.00013590424,0.00029992047,0.0005907206,0.038133148],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9774278,0.015407392,0.0007936792,0.0017627815,0.003453896,0.0011545114],"domain_scores_gemma":[0.97116405,0.016896203,0.00085330853,0.004294584,0.004478769,0.0023131159],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017848054,0.0007511599,0.00044307,0.0015008807,0.0053176354,0.0050266124,0.0018252062,0.0033212863,0.0075170384],"category_scores_gemma":[0.026938098,0.00061913015,0.0006249904,0.0019314407,0.005033302,0.00630844,0.0051430743,0.005316379,0.0027382504],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037598354,0.0041815853,0.018998021,0.0009416865,0.00010599157,0.01315535,0.24413806,0.0067081987,0.01594526,0.072288685,0.040476713,0.5826845],"study_design_scores_gemma":[0.000115825016,0.0024922593,0.005319279,0.00068837637,0.00007552973,0.019592091,0.10412686,0.007054643,0.01840942,0.030630928,0.8112268,0.00026803062],"about_ca_topic_score_codex":0.0014685056,"about_ca_topic_score_gemma":0.0030012438,"teacher_disagreement_score":0.017848054,"about_ca_system_score_codex":0.0016326543,"about_ca_system_score_gemma":0.0026198681,"threshold_uncertainty_score":0.09439069},"labels":[],"label_agreement":null},{"id":"W2097631008","doi":"10.1109/icpc.2008.42","title":"Scenario-Based Comparison of Clone Detection Techniques","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Computer science; Set (abstract data type); Software; Data mining; Artificial intelligence; Operating system; Programming language; Biology","score_opus":0.034191984801920365,"score_gpt":0.2992230280167789,"score_spread":0.26503104321485854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097631008","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5416732,0.0054490482,0.43719366,0.00045632114,0.00017892839,0.00083632115,0.0013742571,0.005162823,0.0076754657],"genre_scores_gemma":[0.81949836,0.001305485,0.17595625,0.000064587344,0.00003561889,0.0002661893,0.0020848792,0.00014811456,0.0006405279],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9845489,0.007837726,0.0011052823,0.0014540754,0.0046411264,0.0004128915],"domain_scores_gemma":[0.881155,0.093195766,0.004438523,0.008338566,0.011779157,0.0010929883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012461946,0.0009704531,0.000979816,0.006928395,0.0005905553,0.0016887336,0.0020810885,0.0018223857,0.0012202186],"category_scores_gemma":[0.07999863,0.00033740012,0.0010577268,0.0038441522,0.00057597103,0.003013588,0.0016095166,0.00088778464,0.0004749183],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036145265,0.0008145147,0.0496449,0.0019349325,0.0015230937,0.00054890336,0.0010151905,0.31301713,0.022001678,0.010997704,0.0046015847,0.59028584],"study_design_scores_gemma":[0.0001996703,0.0021717087,0.033555415,0.00015980039,0.0004421101,0.0018828997,0.00080012553,0.91736615,0.026420934,0.009468014,0.0072860257,0.000247173],"about_ca_topic_score_codex":0.0015778383,"about_ca_topic_score_gemma":0.0015247072,"teacher_disagreement_score":0.012461946,"about_ca_system_score_codex":0.0010707321,"about_ca_system_score_gemma":0.0009185851,"threshold_uncertainty_score":0.06590587},"labels":[],"label_agreement":null},{"id":"W2097666838","doi":"10.1109/iri.2009.5211557","title":"Scalability improvement in software evaluation methodologies","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Scalability; Computer science; Software; Dependency (UML); Adaptation (eye); Software engineering; Software metric; Software development; Software quality; Reliability engineering; Engineering; Database; Operating system","score_opus":0.08660750563061385,"score_gpt":0.38014303471656613,"score_spread":0.29353552908595226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097666838","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028597107,0.0067267455,0.9434829,0.0030116003,0.00020814333,0.0010724242,0.00009925385,0.0017658309,0.01503609],"genre_scores_gemma":[0.2314546,0.002448459,0.76219654,0.00050796836,0.00022686229,0.0008830067,0.00024625246,0.00036040694,0.0016759128],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.88854164,0.0568005,0.010400964,0.0054104547,0.03705332,0.0017931511],"domain_scores_gemma":[0.8114328,0.10117138,0.014016909,0.021884982,0.050103296,0.0013907445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0829753,0.0018849586,0.0015549131,0.0075917267,0.0014006423,0.0051967413,0.0031461976,0.001802407,0.001461],"category_scores_gemma":[0.18572311,0.001206118,0.0017423963,0.0044259736,0.0019650345,0.008692028,0.0046614064,0.003295519,0.00066062436],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018481982,0.0004109728,0.010324247,0.0030216402,0.0004561997,0.00034823618,0.0021764524,0.047770668,0.011776096,0.09602508,0.0060480777,0.82145756],"study_design_scores_gemma":[0.00034651306,0.0014574893,0.016653981,0.0049256464,0.0008713731,0.0019239036,0.0027660741,0.5379247,0.030137429,0.28515625,0.11738946,0.00044721476],"about_ca_topic_score_codex":0.0022430953,"about_ca_topic_score_gemma":0.0019283822,"teacher_disagreement_score":0.0829753,"about_ca_system_score_codex":0.0041329083,"about_ca_system_score_gemma":0.0049041654,"threshold_uncertainty_score":0.43882054},"labels":[],"label_agreement":null},{"id":"W2097710325","doi":"10.1145/1082983.1083180","title":"A design for evidence - based soft research","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Software engineering; Process (computing); Architecture; Data science; Empirical evidence; Triangulation; Management science; Systems engineering; Engineering","score_opus":0.13352022884253242,"score_gpt":0.34857451112989774,"score_spread":0.2150542822873653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097710325","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0063235145,0.005373621,0.5565718,0.028638178,0.006633097,0.33110586,0.004998325,0.0013417837,0.059013914],"genre_scores_gemma":[0.016313855,0.0011963482,0.68506753,0.0041268226,0.00022910946,0.28906745,0.0005270271,0.00014689066,0.0033251042],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.5863895,0.3318872,0.041233666,0.016864074,0.019807007,0.003818595],"domain_scores_gemma":[0.54435056,0.31107935,0.017527537,0.06402895,0.051847897,0.011165677],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.27211103,0.0018076521,0.005071568,0.019785011,0.006555756,0.018687122,0.007244592,0.012632221,0.05357883],"category_scores_gemma":[0.35290664,0.0034315556,0.006171012,0.016913207,0.009348892,0.016589256,0.01778911,0.010463446,0.011434088],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043335357,0.0010394496,0.0024504124,0.04109932,0.0010308723,0.00096836826,0.012207229,0.0023624788,0.0022084406,0.5513555,0.036618832,0.34432548],"study_design_scores_gemma":[0.00713949,0.003331162,0.001991999,0.03209453,0.0016496411,0.0006276495,0.009093968,0.003404295,0.002885615,0.26613382,0.67132694,0.00032092596],"about_ca_topic_score_codex":0.0011334178,"about_ca_topic_score_gemma":0.0015858933,"teacher_disagreement_score":0.72788894,"about_ca_system_score_codex":0.013042513,"about_ca_system_score_gemma":0.034387667,"threshold_uncertainty_score":0.8976167},"labels":[],"label_agreement":null},{"id":"W2097737237","doi":"10.5555/2820282.2820314","title":"License usage and changes: a large-scale study of Java projects on GitHub","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"License; Commit; Computer science; TRACE (psycholinguistics); Java; Artifact (error); Reuse; Software; Software deployment; Software development; Software engineering; Point (geometry); Software evolution; Grounded theory; Data science; Computer security; Qualitative research; Database; Software construction; Programming language; Engineering; Operating system","score_opus":0.0592872006152315,"score_gpt":0.3000912609747387,"score_spread":0.2408040603595072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097737237","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9987619,0.00006141572,0.00027410983,0.00010602663,0.000003617908,0.000033996777,0.0001501912,0.000008301329,0.0006003776],"genre_scores_gemma":[0.9985135,0.000094963296,0.00039699065,0.00006639497,0.0000084239555,0.00008473682,0.0003168089,0.000036240366,0.00048186394],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9905578,0.0043798666,0.00060332735,0.0012235533,0.0022837752,0.0009517709],"domain_scores_gemma":[0.9102152,0.056600202,0.017842632,0.003691942,0.0077412445,0.00390878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007960342,0.00031103424,0.0005593246,0.00568159,0.002670861,0.0034078392,0.0014128872,0.0011236024,0.0017091435],"category_scores_gemma":[0.04901208,0.00058244437,0.00031062265,0.008474606,0.0034434437,0.0039122845,0.004178909,0.0019284374,0.0005180933],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013943299,0.00072764576,0.6929827,0.00030392554,0.000069201036,0.0017414892,0.2751378,0.00033646912,0.0014007234,0.00083665614,0.0015692399,0.024754765],"study_design_scores_gemma":[0.00000845063,0.00013390745,0.83405936,0.00011941355,0.0000127248895,0.00027187922,0.16073018,0.0007959791,0.00037341812,0.00021676929,0.0032334588,0.00004434203],"about_ca_topic_score_codex":0.027346395,"about_ca_topic_score_gemma":0.053648707,"teacher_disagreement_score":0.027346395,"about_ca_system_score_codex":0.0033670396,"about_ca_system_score_gemma":0.0021661331,"threshold_uncertainty_score":0.054374456},"labels":[],"label_agreement":null},{"id":"W2097750323","doi":"10.1109/tse.2008.26","title":"Asking and Answering Questions during a Programming Change Task","year":2008,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":314,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; University of Alberta; Vrije Universiteit Brussel; University of Washington","keywords":"Computer science; Programmer; Ask price; Task (project management); Key (lock); Programming language; Question answering; Code (set theory); Software engineering; World Wide Web; Data science; Information retrieval","score_opus":0.022144955307222974,"score_gpt":0.23619539299239536,"score_spread":0.2140504376851724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097750323","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9550355,0.00032917765,0.036767155,0.0016592955,0.00003826265,0.00046819478,0.00028252616,0.0010382457,0.0043815514],"genre_scores_gemma":[0.9519398,0.00026437122,0.043155007,0.00090824364,0.000048127895,0.00051241874,0.000560074,0.00028482915,0.0023269735],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9675163,0.023655517,0.001625967,0.002573856,0.0031507933,0.0014775818],"domain_scores_gemma":[0.6370504,0.32284904,0.013663167,0.012380022,0.009353153,0.0047043124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019963067,0.000797919,0.0010076755,0.0022958203,0.0026854621,0.0032947452,0.0021485249,0.0041041854,0.0029812397],"category_scores_gemma":[0.18126258,0.0011956067,0.0006895764,0.0014714572,0.0021204131,0.0055202525,0.003897749,0.0032844313,0.0010571798],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012317931,0.00076429924,0.18357284,0.0011807234,0.00013655235,0.0018620827,0.62082964,0.0015760697,0.044308584,0.0033806784,0.006171025,0.13498572],"study_design_scores_gemma":[0.00045278118,0.0037935546,0.30222824,0.0010819273,0.00034786077,0.0060992027,0.44809082,0.050607804,0.035533577,0.016047873,0.13485453,0.00086195825],"about_ca_topic_score_codex":0.0030734115,"about_ca_topic_score_gemma":0.0032775623,"teacher_disagreement_score":0.019963067,"about_ca_system_score_codex":0.0011955246,"about_ca_system_score_gemma":0.0016340293,"threshold_uncertainty_score":0.1055761},"labels":[],"label_agreement":null},{"id":"W2097886687","doi":"10.1109/wcre.2010.37","title":"A Case Study of Bias in Bug-Fix Datasets","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Traceability; Generality; Heuristics; Software quality; Quality (philosophy); Software; Source code; Data quality; Code (set theory); Software bug; Data mining; Data science; Software engineering; Software development; Programming language","score_opus":0.06995749814094142,"score_gpt":0.3397108915861387,"score_spread":0.2697533934451973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097886687","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9306396,0.0017206203,0.05448679,0.0037393216,0.000117158495,0.0003974238,0.005372236,0.0010938587,0.0024328888],"genre_scores_gemma":[0.9368808,0.00023886196,0.05650394,0.0006572169,0.000082995735,0.00021329476,0.0047886246,0.00022532973,0.00040885963],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.88419515,0.08351575,0.0068716262,0.011667802,0.012156459,0.0015931341],"domain_scores_gemma":[0.38215452,0.5166614,0.01989371,0.06023876,0.01904623,0.002005461],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.10536785,0.00084038475,0.001134926,0.0031936439,0.0021455674,0.0027011812,0.0023237958,0.003687972,0.00086697005],"category_scores_gemma":[0.36834288,0.0006034344,0.0014113319,0.0076027242,0.0035082162,0.0039629703,0.0027158726,0.002771399,0.00036754628],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003047767,0.0014860205,0.68572485,0.0017208984,0.0019591223,0.0024298131,0.010092034,0.10247559,0.006474254,0.037440594,0.028479004,0.11867003],"study_design_scores_gemma":[0.0020222387,0.0016499319,0.29660684,0.0012177315,0.0010753712,0.007042351,0.0062221684,0.44665283,0.0334118,0.13426676,0.069198415,0.00063354056],"about_ca_topic_score_codex":0.009543196,"about_ca_topic_score_gemma":0.011192571,"teacher_disagreement_score":0.89463216,"about_ca_system_score_codex":0.003087745,"about_ca_system_score_gemma":0.0023318245,"threshold_uncertainty_score":0.55724514},"labels":[],"label_agreement":null},{"id":"W2097906042","doi":"10.1109/icsm.2002.1167793","title":"Measuring software functional size: towards an effective measurement of complexity","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Programming complexity; Simple (philosophy); Worst-case complexity; Software; Computational complexity theory; Theoretical computer science; Software system; Algorithm; Software construction; Programming language","score_opus":0.09861013940244175,"score_gpt":0.2735393058588634,"score_spread":0.17492916645642165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097906042","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023392538,0.0008476932,0.9721016,0.0005931345,0.00007057669,0.00013118483,0.000107467786,0.0003737707,0.002382085],"genre_scores_gemma":[0.27599356,0.0012394595,0.72037375,0.00019182848,0.00020123248,0.00062600186,0.00028734375,0.00023178136,0.00085513486],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9768722,0.0072994633,0.0016697526,0.0021821754,0.011458675,0.0005176274],"domain_scores_gemma":[0.89908385,0.06336906,0.012341017,0.011722097,0.012027972,0.0014560141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012442873,0.0022827825,0.001344706,0.008733632,0.000982856,0.0044165053,0.0030360795,0.0026741363,0.0014896997],"category_scores_gemma":[0.096199796,0.0010070638,0.0009619334,0.0058148853,0.0050465656,0.01576216,0.0043239924,0.002484326,0.0005963734],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025253696,0.0003520352,0.034994278,0.0025362477,0.00029855387,0.00018821604,0.00319679,0.043610215,0.048773225,0.27778476,0.0038755697,0.5841376],"study_design_scores_gemma":[0.000084241205,0.0016000082,0.044024732,0.00072681054,0.00022232438,0.001012234,0.0020046725,0.34991008,0.07093629,0.50202906,0.026893275,0.0005563081],"about_ca_topic_score_codex":0.0007278957,"about_ca_topic_score_gemma":0.0006150731,"teacher_disagreement_score":0.012442873,"about_ca_system_score_codex":0.0014307179,"about_ca_system_score_gemma":0.0013685505,"threshold_uncertainty_score":0.06580496},"labels":[],"label_agreement":null},{"id":"W2097956891","doi":"10.1109/ccece.1999.807208","title":"Validation of software effort models","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Data modeling; Software quality; Software; Data mining; Set (abstract data type); Data set; Software metric; Software development; Artificial intelligence; Software engineering; Programming language","score_opus":0.026395172069029905,"score_gpt":0.26150583370093466,"score_spread":0.23511066163190475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097956891","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8897399,0.00036136218,0.10222373,0.00033545442,0.000095841526,0.00029355378,0.0023656206,0.0008798948,0.0037045996],"genre_scores_gemma":[0.9755136,0.00012699005,0.020023229,0.000058475987,0.000015614987,0.00031979324,0.002790984,0.00013963219,0.0010117234],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9868687,0.006570364,0.0010017295,0.0021092454,0.0028429944,0.0006069119],"domain_scores_gemma":[0.8984587,0.07576587,0.0031144696,0.010865597,0.011245561,0.0005498107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025827985,0.0018583102,0.0013049323,0.0037415845,0.00077186484,0.0022256868,0.0033977684,0.0020154612,0.0018078897],"category_scores_gemma":[0.08823493,0.00066674716,0.002410693,0.0026040433,0.0011209596,0.0027114863,0.0014828242,0.0020208845,0.00079182687],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009934902,0.000767594,0.045626223,0.00030814967,0.0005527412,0.0001629305,0.00059263734,0.8921848,0.0021288553,0.00608154,0.0017273831,0.048873655],"study_design_scores_gemma":[0.00009153447,0.0007507376,0.020932818,0.00009818642,0.00007973225,0.00007869827,0.00019361575,0.9691056,0.00307548,0.004356315,0.0011802096,0.000057137157],"about_ca_topic_score_codex":0.016009906,"about_ca_topic_score_gemma":0.009623227,"teacher_disagreement_score":0.025827985,"about_ca_system_score_codex":0.0036767721,"about_ca_system_score_gemma":0.0023247458,"threshold_uncertainty_score":0.1365931},"labels":[],"label_agreement":null},{"id":"W2098007560","doi":"10.5539/cis.v8n4p51","title":"A Smart Algorithm for USE-Cases Production Based on Name Entity Recognition","year":2015,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Unified Modeling Language; Use Case Diagram; Process (computing); Software; Set (abstract data type); Finite-state machine; Natural language; Programming language; Plain text; Natural language processing; Software engineering; Class diagram; Algorithm; Artificial intelligence","score_opus":0.06687359796104474,"score_gpt":0.2847958421025112,"score_spread":0.21792224414146644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098007560","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019400971,0.000046868703,0.96779823,0.000084853535,0.00004839034,0.00051092816,0.00033915238,0.02789691,0.0013345408],"genre_scores_gemma":[0.016903564,0.000049217233,0.97886056,0.000056983597,0.000017243803,0.00045912777,0.0009689077,0.0009339928,0.0017504559],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99580044,0.00074791646,0.0006344622,0.0015183735,0.0011148426,0.00018397743],"domain_scores_gemma":[0.99563307,0.0020336157,0.00039830897,0.0007439936,0.0010630479,0.00012793278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031617668,0.0021047576,0.0012194779,0.0049676443,0.0013406632,0.0030801962,0.0024854974,0.0017095642,0.015709946],"category_scores_gemma":[0.010830472,0.0010335225,0.0014713536,0.0023225371,0.0013718852,0.002735952,0.0018845334,0.0019712914,0.012991593],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030551205,0.00024071542,0.0025074848,0.0005149064,0.00007693273,0.00036817818,0.0006950323,0.0054887254,0.036373936,0.013278422,0.020054715,0.9200955],"study_design_scores_gemma":[0.00025544793,0.00035195207,0.0045365607,0.00025284052,0.00016973315,0.0022676804,0.00054600247,0.58858716,0.21866208,0.026526574,0.1575703,0.0002736178],"about_ca_topic_score_codex":0.0021372063,"about_ca_topic_score_gemma":0.0021081294,"teacher_disagreement_score":0.015709946,"about_ca_system_score_codex":0.00093989726,"about_ca_system_score_gemma":0.002088149,"threshold_uncertainty_score":0.052554965},"labels":[],"label_agreement":null},{"id":"W2098046929","doi":"10.1002/spe.1053","title":"Prioritizing the creation of unit tests in legacy software systems","year":2011,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada); Queen's University","funders":"","keywords":"Unit testing; Test-driven development; Computer science; Software engineering; Legacy system; Adaptation (eye); Software development; Development testing; Heuristics; Software; Software development process; Systems engineering; Engineering; Operating system","score_opus":0.045011511546345535,"score_gpt":0.3163417384664236,"score_spread":0.2713302269200781,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098046929","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8520335,0.0010004711,0.13714598,0.00088666804,0.0000657261,0.00049764354,0.00009836611,0.0035680367,0.004703625],"genre_scores_gemma":[0.78748167,0.00025895934,0.21032223,0.00016068928,0.000030524352,0.00013684305,0.0002754377,0.00033130948,0.0010023789],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99255425,0.004291069,0.0006072571,0.0006283553,0.0015801234,0.00033906958],"domain_scores_gemma":[0.9260782,0.052133963,0.007034401,0.005183284,0.007465438,0.0021047387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008150267,0.0006879695,0.0005132334,0.0025981637,0.0006389734,0.0020441583,0.0017128697,0.00062884245,0.0014493695],"category_scores_gemma":[0.0499706,0.0006889839,0.0004082618,0.0010660365,0.00048010473,0.0014471186,0.0012833611,0.00091080734,0.00049583503],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073826505,0.0007812963,0.06846925,0.00058979087,0.000074936805,0.0007935425,0.0017070005,0.036252405,0.03279301,0.0025676435,0.0033889888,0.85184395],"study_design_scores_gemma":[0.00057572144,0.0038381226,0.109691866,0.0005839913,0.00038117042,0.0033551862,0.0027408462,0.6874912,0.14561006,0.01867835,0.026771715,0.00028176862],"about_ca_topic_score_codex":0.0026120935,"about_ca_topic_score_gemma":0.0046614376,"teacher_disagreement_score":0.008150267,"about_ca_system_score_codex":0.0012896426,"about_ca_system_score_gemma":0.0024630385,"threshold_uncertainty_score":0.043103218},"labels":[],"label_agreement":null},{"id":"W2098092920","doi":"10.1109/fuzzy.2005.1452434","title":"Fuzzy Structural Dependency Constraints in Software Release Planning","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Fuzzy logic; Dependency (UML); Software; Software development; Software release life cycle; Identification (biology); Software engineering; Data mining; Software construction; Artificial intelligence; Programming language","score_opus":0.017928318302665707,"score_gpt":0.2819274696765192,"score_spread":0.2639991513738535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098092920","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06744037,0.001365427,0.91911054,0.0010526865,0.00005440057,0.00014006982,0.00030497,0.000119377066,0.0104121445],"genre_scores_gemma":[0.7952139,0.001335648,0.19974789,0.0001356153,0.00010903876,0.0002732777,0.000418861,0.000045126635,0.0027207073],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973053,0.0010728893,0.00017761734,0.00036113537,0.00090314465,0.00017984358],"domain_scores_gemma":[0.98735726,0.010656472,0.00077887875,0.0002312338,0.0007922968,0.00018390786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004037324,0.0008116637,0.0009784717,0.0018267541,0.00095228665,0.0017560091,0.0011369995,0.0011558045,0.002114721],"category_scores_gemma":[0.012115729,0.000805337,0.0010182371,0.0021302365,0.0017449736,0.0029691136,0.0010215776,0.0014158464,0.00021140427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000107070526,0.00005129692,0.0008527371,0.00015963386,0.000058400634,0.0004654508,0.00034918703,0.8411489,0.0010136196,0.12352053,0.0010633857,0.031209817],"study_design_scores_gemma":[0.000027096414,0.000043622247,0.00052213523,0.0000400699,0.000027027148,0.00007398369,0.00008540748,0.8422069,0.00066484686,0.15500854,0.0012667588,0.000033653345],"about_ca_topic_score_codex":0.010691851,"about_ca_topic_score_gemma":0.008168961,"teacher_disagreement_score":0.010691851,"about_ca_system_score_codex":0.002345403,"about_ca_system_score_gemma":0.0014652405,"threshold_uncertainty_score":0.021351635},"labels":[],"label_agreement":null},{"id":"W2098124758","doi":"10.1109/ccece.2013.6567821","title":"Near-miss clone patterns in web applications: An empirical study with industrial systems","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Japan Society for the Promotion of Science","keywords":"JavaScript; Computer science; Source code; Web application; Programming language; World Wide Web; Cloning (programming); Dynamic web page; Code (set theory); Open source; Software maintenance; Web modeling; Software engineering; Software; Web page; Software system","score_opus":0.04884154400166066,"score_gpt":0.3140051924206891,"score_spread":0.2651636484190284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098124758","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9994795,0.00004340888,0.00024351255,0.000019379373,5.2945984e-7,0.000023064396,0.000014305517,0.0000045541533,0.0001718032],"genre_scores_gemma":[0.9986859,0.00009286456,0.00084599195,0.000025489335,0.0000032155033,0.000034685174,0.000086817556,0.000009453612,0.00021558674],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9947871,0.002279946,0.0004778913,0.00075028784,0.0014022297,0.00030256604],"domain_scores_gemma":[0.9029105,0.06847024,0.015125561,0.0045062904,0.0069574495,0.0020300339],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057240855,0.00034598287,0.00041282666,0.0021507354,0.0012917104,0.0013272421,0.0010707736,0.0011805131,0.0009030961],"category_scores_gemma":[0.047925346,0.00040951846,0.0004045519,0.002258045,0.001616546,0.0026607572,0.0014736693,0.0012515188,0.00033250503],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023484284,0.0034797855,0.93461365,0.00021429025,0.00008676731,0.0016201375,0.029593058,0.0012102546,0.0016864124,0.0003269659,0.00041720443,0.026516613],"study_design_scores_gemma":[0.000053291402,0.002271844,0.9493882,0.00009955328,0.0000842721,0.002385844,0.029273933,0.011938729,0.0021300437,0.0005119413,0.0018117116,0.000050667837],"about_ca_topic_score_codex":0.003468445,"about_ca_topic_score_gemma":0.0056559835,"teacher_disagreement_score":0.0057240855,"about_ca_system_score_codex":0.0008493986,"about_ca_system_score_gemma":0.000748648,"threshold_uncertainty_score":0.030272245},"labels":[],"label_agreement":null},{"id":"W2098302578","doi":"10.1109/wcre.2003.1287270","title":"First international workshop on refactoring : achievements, challenges, and effects (REFACE'03)","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code refactoring; Restructuring; Computer science; Software engineering; Scalability; Object-oriented programming; Position paper; Formalism (music); Programming language; Software; World Wide Web; Political science","score_opus":0.03213658356081325,"score_gpt":0.2737441751458723,"score_spread":0.24160759158505904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098302578","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03843384,0.089932814,0.62616354,0.06650556,0.05564873,0.0014083231,0.003193517,0.010283329,0.10843026],"genre_scores_gemma":[0.11227283,0.053924885,0.44591412,0.0140012745,0.0102939475,0.0013989339,0.019734297,0.0062595145,0.33620015],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9922621,0.0032174822,0.00042965924,0.0010852045,0.0021318952,0.000873667],"domain_scores_gemma":[0.98037624,0.007088741,0.00041145875,0.0025095437,0.0062472094,0.0033667856],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017547585,0.0020924648,0.0025831233,0.0020650015,0.0019853418,0.008509364,0.004644534,0.0058423197,0.031454317],"category_scores_gemma":[0.018043257,0.0011138375,0.002405337,0.0030245276,0.001665544,0.00806399,0.0053768535,0.0071784654,0.012696387],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062985625,0.0008938712,0.0011809868,0.00080094923,0.000099269964,0.00042249204,0.0018887438,0.0029502055,0.005028788,0.013151229,0.4294159,0.54353774],"study_design_scores_gemma":[0.0002388324,0.0005780878,0.0023594925,0.0008270237,0.000109423374,0.00071792514,0.0016922025,0.008018778,0.0049960315,0.019357817,0.96098566,0.00011874443],"about_ca_topic_score_codex":0.008052849,"about_ca_topic_score_gemma":0.011662015,"teacher_disagreement_score":0.031454317,"about_ca_system_score_codex":0.0023588492,"about_ca_system_score_gemma":0.0048079775,"threshold_uncertainty_score":0.105225205},"labels":[],"label_agreement":null},{"id":"W2098347799","doi":"10.1002/smr.416","title":"Near‐miss function clones in open source software: an empirical study","year":2009,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Saskatchewan","funders":"","keywords":"clone (Java method); Cloning (programming); Source code; Java; Computer science; Open source; Linux kernel; Benchmark (surveying); Operating system; Function (biology); Software maintenance; Software; Software system; Programming language; Biology; Genetics; Gene","score_opus":0.06872976237691421,"score_gpt":0.40221598404488773,"score_spread":0.33348622166797354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098347799","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9991726,0.00007800205,0.00047968052,0.00001947997,8.1961196e-7,0.0000124441485,0.000047121866,0.000008817779,0.00018093728],"genre_scores_gemma":[0.9990056,0.000059079874,0.00059360365,0.000012798333,0.0000019734332,0.00001585823,0.00014356978,0.000009671244,0.00015782424],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9891452,0.0042904653,0.0010622876,0.0016105794,0.0035397639,0.00035159622],"domain_scores_gemma":[0.79712623,0.14448878,0.02933967,0.009771983,0.01732925,0.0019441303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007910067,0.00028569356,0.00037001885,0.002611527,0.00073366135,0.0012059628,0.0010771151,0.0006994308,0.0010233285],"category_scores_gemma":[0.078840025,0.0003344804,0.00031458787,0.0023435515,0.0015121772,0.0021469053,0.0015005528,0.000892326,0.0002989592],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017404277,0.00029713922,0.96985424,0.0001187231,0.000070775386,0.00053087645,0.0072593875,0.0004315709,0.0010975426,0.00026134515,0.00033497016,0.01956939],"study_design_scores_gemma":[0.000019870251,0.0005839092,0.9780352,0.000076988705,0.00007374757,0.0026743931,0.008191163,0.0059157615,0.002266383,0.00027490637,0.0018538951,0.00003363166],"about_ca_topic_score_codex":0.0024053564,"about_ca_topic_score_gemma":0.0037008685,"teacher_disagreement_score":0.007910067,"about_ca_system_score_codex":0.0006359254,"about_ca_system_score_gemma":0.0005050674,"threshold_uncertainty_score":0.041832983},"labels":[],"label_agreement":null},{"id":"W2098349183","doi":"10.1109/icsm.1995.526550","title":"A sizing measure for adaptive maintenance work products","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Sizing; Function point; Granularity; Computer science; Measure (data warehouse); Work (physics); Process (computing); Product (mathematics); Reliability engineering; Function (biology); Software; Perspective (graphical); Interval (graph theory); Productivity; Industrial engineering; Software development; Engineering; Data mining; Mathematics; Artificial intelligence; Mechanical engineering","score_opus":0.05726241815363975,"score_gpt":0.24150781486305065,"score_spread":0.1842453967094109,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098349183","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15159996,0.00054862135,0.83843386,0.00010735743,0.00007792392,0.00029477986,0.00053038116,0.0011203786,0.0072866655],"genre_scores_gemma":[0.46143892,0.000186424,0.535742,0.00003486345,0.000048874193,0.00023602427,0.0006857754,0.000116425486,0.0015106697],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9968706,0.00040251078,0.00033214406,0.00045923344,0.0018662519,0.000069322836],"domain_scores_gemma":[0.9850307,0.007379554,0.0025063588,0.001942666,0.002946697,0.00019416325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023155967,0.00070659816,0.00045238354,0.0035541402,0.00029249868,0.0012906472,0.00087466644,0.0005971361,0.0020103154],"category_scores_gemma":[0.01539799,0.00025920544,0.0005000181,0.0027037773,0.00065276056,0.0017147453,0.000544813,0.0006347171,0.00059299247],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005719032,0.00026458228,0.032825384,0.00065886346,0.000117988166,0.0001182212,0.0007560768,0.048913486,0.11469954,0.019714374,0.0023258699,0.77903366],"study_design_scores_gemma":[0.00014480662,0.0048539457,0.19055429,0.00022727768,0.000329881,0.0016936127,0.0010411342,0.54084295,0.19774993,0.027535656,0.03471357,0.0003128933],"about_ca_topic_score_codex":0.00071608153,"about_ca_topic_score_gemma":0.0008264045,"teacher_disagreement_score":0.0035541402,"about_ca_system_score_codex":0.00070659106,"about_ca_system_score_gemma":0.0003780127,"threshold_uncertainty_score":0.0122461915},"labels":[],"label_agreement":null},{"id":"W2098350227","doi":"10.1145/985921.986081","title":"Recent developments in text-entry error rate measurement","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Word error rate; Text entry; Key (lock); Character (mathematics); Software; Work (physics); Error analysis; Error detection and correction; Theoretical computer science; Algorithm; Artificial intelligence; Programming language; Computer security; Human–computer interaction; Mathematics; Engineering","score_opus":0.05738068243037023,"score_gpt":0.2807148201110655,"score_spread":0.2233341376806953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098350227","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051554227,0.0060096066,0.98187596,0.0005143317,0.00036139227,0.00017874628,0.00039455396,0.0029426345,0.0025674426],"genre_scores_gemma":[0.06382838,0.0063161515,0.9201826,0.00042998328,0.0016365579,0.00053626206,0.0012309351,0.0014562428,0.0043828],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.93275285,0.020523114,0.008199672,0.011781522,0.025845386,0.00089746655],"domain_scores_gemma":[0.69864726,0.15602818,0.024553804,0.035432175,0.08240679,0.0029317902],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024669547,0.0037222533,0.0037815394,0.0125545645,0.0012127415,0.0067703472,0.0064827204,0.0031671093,0.0043065306],"category_scores_gemma":[0.16242398,0.0018965742,0.0015381521,0.011488627,0.0029133789,0.010125755,0.0036424203,0.005477679,0.0043279044],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005122205,0.00027810194,0.0072699776,0.0019918142,0.000249753,0.00011024397,0.0007056451,0.014898564,0.02178125,0.038057465,0.0067454833,0.9073995],"study_design_scores_gemma":[0.0001931575,0.0030433938,0.027887074,0.0024643121,0.0010174187,0.0040450906,0.0009865825,0.40817636,0.22091325,0.084381916,0.24484491,0.0020465464],"about_ca_topic_score_codex":0.0026761885,"about_ca_topic_score_gemma":0.002022573,"teacher_disagreement_score":0.024669547,"about_ca_system_score_codex":0.0024921694,"about_ca_system_score_gemma":0.0024117737,"threshold_uncertainty_score":0.13046658},"labels":[],"label_agreement":null},{"id":"W2098427431","doi":"10.1109/ccece.2003.1226096","title":"Discriminatory software metric selection via a grid of interconnected multilayer perceptrons","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Institute for Biodiagnostics","funders":"","keywords":"Computer science; Perceptron; Classifier (UML); Heuristics; Data mining; Software quality; Grid; Software; Artificial intelligence; Metric (unit); Machine learning; Source code; Feature selection; Software metric; Rule of thumb; Pattern recognition (psychology); Software development; Artificial neural network; Algorithm; Engineering; Mathematics","score_opus":0.014410278427675975,"score_gpt":0.2593157559667531,"score_spread":0.24490547753907713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098427431","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19138291,0.00012480101,0.805637,0.0001750759,0.000025210364,0.00008440023,0.000062370134,0.0009668348,0.0015414513],"genre_scores_gemma":[0.87061465,0.000043913977,0.12829657,0.000052512456,0.000011148839,0.00010835405,0.00009173055,0.00003534461,0.0007458349],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984725,0.0005683281,0.0001247749,0.00031396715,0.00033335047,0.00018710218],"domain_scores_gemma":[0.9961139,0.0019176952,0.00047793466,0.00053850963,0.00080877694,0.00014309796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003121305,0.0007539744,0.0010142004,0.0009697035,0.00045378224,0.0012893878,0.0011171567,0.0007908504,0.0006846798],"category_scores_gemma":[0.011499897,0.00050488213,0.000701644,0.00085112307,0.0009236932,0.0014433613,0.001245807,0.0011299751,0.00027023352],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047499946,0.00014017022,0.009018197,0.00009204684,0.00012276172,0.0001431782,0.00025521062,0.7515878,0.013943885,0.0067059617,0.00082935195,0.2166864],"study_design_scores_gemma":[0.000012768739,0.000063203064,0.000837672,0.000006644562,0.000013037317,0.000020562145,0.000020369487,0.9912582,0.0027429154,0.004838997,0.00017796403,0.000007664784],"about_ca_topic_score_codex":0.0024981971,"about_ca_topic_score_gemma":0.00204179,"teacher_disagreement_score":0.003121305,"about_ca_system_score_codex":0.0009941241,"about_ca_system_score_gemma":0.0005677119,"threshold_uncertainty_score":0.016507208},"labels":[],"label_agreement":null},{"id":"W2098673848","doi":"10.1145/2025113.2025155","title":"High-impact defects","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Surprise; Computer science; Software bug; Software; Guard (computer science); Operating system","score_opus":0.03303628079213807,"score_gpt":0.2665415618860166,"score_spread":0.2335052810938785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098673848","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9117868,0.0012668632,0.059524577,0.0005891197,0.00018170226,0.0005423311,0.004960876,0.0028294711,0.018318165],"genre_scores_gemma":[0.98394275,0.00035175966,0.008625834,0.00011062235,0.000029576577,0.00009302978,0.0029689195,0.00014473684,0.0037328515],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99612254,0.00035339623,0.00025825546,0.00072141556,0.002105411,0.00043896082],"domain_scores_gemma":[0.92970484,0.040917017,0.015191761,0.004704658,0.008238228,0.0012435948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024889957,0.0008686487,0.00045631416,0.00356081,0.00045398387,0.0016295151,0.0011865604,0.0010776335,0.0074635074],"category_scores_gemma":[0.035138857,0.00041094588,0.0009112291,0.0022452595,0.0007851847,0.0020232978,0.0013774065,0.0012575579,0.0013043594],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045229876,0.00081749295,0.7563042,0.00091140467,0.00025081262,0.0017584489,0.00067352416,0.04323078,0.0058465395,0.008121073,0.008652071,0.17298138],"study_design_scores_gemma":[0.0000673042,0.0009783766,0.803194,0.00038562992,0.0002895528,0.0040927115,0.000807506,0.14618866,0.011969072,0.014818284,0.017054377,0.00015448515],"about_ca_topic_score_codex":0.0036819635,"about_ca_topic_score_gemma":0.0049046515,"teacher_disagreement_score":0.0074635074,"about_ca_system_score_codex":0.0012053732,"about_ca_system_score_gemma":0.0008173664,"threshold_uncertainty_score":0.024967968},"labels":[],"label_agreement":null},{"id":"W2098711396","doi":"","title":"ISO/IEC SQuaRE. The second generation of standards for software product quality","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Quality (philosophy); Software; Product (mathematics); Software quality control; Software quality; Engineering; Computer science; Systems engineering; Software engineering; Manufacturing engineering; Software development; Operating system; Mathematics","score_opus":0.05594303300798379,"score_gpt":0.33425110349062653,"score_spread":0.2783080704826427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098711396","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0122625325,0.047123734,0.31341407,0.026111066,0.01749359,0.0031229611,0.022945724,0.010345586,0.5471808],"genre_scores_gemma":[0.09655056,0.038261738,0.50763124,0.010353707,0.0026578533,0.0054388857,0.06966428,0.0043173097,0.26512444],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9822939,0.0020772808,0.0017124083,0.0008211909,0.01261645,0.00047873126],"domain_scores_gemma":[0.98302186,0.001976217,0.0011466725,0.0015697492,0.01187589,0.00040959317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006060244,0.0020333447,0.000873654,0.0062735016,0.00108904,0.003263347,0.0020349526,0.0038954385,0.011113788],"category_scores_gemma":[0.017805574,0.0007145917,0.00094975263,0.0057723704,0.0015380454,0.004067987,0.001671888,0.0038016634,0.014824062],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028164408,0.0002681146,0.001762857,0.0022151861,0.00004373717,0.00038471393,0.00073281094,0.0022105437,0.0078080813,0.109218255,0.45661134,0.4184627],"study_design_scores_gemma":[0.00003804738,0.00009861852,0.0022289252,0.00047694417,0.00002819788,0.0003069031,0.00014671376,0.00090551353,0.0020260164,0.013637861,0.980078,0.000028141088],"about_ca_topic_score_codex":0.008668983,"about_ca_topic_score_gemma":0.005050937,"teacher_disagreement_score":0.011113788,"about_ca_system_score_codex":0.0019565057,"about_ca_system_score_gemma":0.009537217,"threshold_uncertainty_score":0.03717935},"labels":[],"label_agreement":null},{"id":"W2098793283","doi":"10.1109/csmr.2000.827305","title":"Design properties and object-oriented software changeability","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Inheritance (genetic algorithm); Object-oriented design; Software; Set (abstract data type); Dependency (UML); Object-oriented programming; Software system; Class (philosophy); Object (grammar); Software design; Software engineering; Software development; Programming language; Artificial intelligence","score_opus":0.07748752429194466,"score_gpt":0.23683594832512203,"score_spread":0.15934842403317737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098793283","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67758423,0.0011925059,0.30693266,0.0014044669,0.000048760212,0.00019499304,0.0001309531,0.0007539346,0.011757405],"genre_scores_gemma":[0.97182155,0.00021387747,0.02708385,0.000074955584,0.00004411242,0.00012472103,0.00013207975,0.00007433908,0.00043055686],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9920947,0.0034673698,0.00070708647,0.00076819083,0.0025821512,0.00038048846],"domain_scores_gemma":[0.8646362,0.08794926,0.02892694,0.009887613,0.0075280634,0.0010719651],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074184015,0.000561312,0.0004059987,0.0033702159,0.0006343615,0.0018860928,0.0005365394,0.0011214382,0.00091253896],"category_scores_gemma":[0.07675853,0.0005080417,0.0006978296,0.001813981,0.0030715198,0.004427348,0.0007812559,0.0010922467,0.00017662333],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036676746,0.0009154577,0.30717444,0.0013481096,0.00039627057,0.0012420778,0.005383804,0.0966162,0.025329338,0.34362936,0.0014383572,0.2161598],"study_design_scores_gemma":[0.00016969016,0.0011091384,0.19979732,0.00042129852,0.00033098547,0.0023104383,0.0016685659,0.2490254,0.022352582,0.50878906,0.013868036,0.0001575353],"about_ca_topic_score_codex":0.0012732102,"about_ca_topic_score_gemma":0.00069943815,"teacher_disagreement_score":0.0074184015,"about_ca_system_score_codex":0.0013266085,"about_ca_system_score_gemma":0.0007637139,"threshold_uncertainty_score":0.03923279},"labels":[],"label_agreement":null},{"id":"W2098969028","doi":"10.5555/857171.857210","title":"The Dangerous 'All' in Specifications","year":2000,"lang":"en","type":"article","venue":"International Workshop on Software Specification and Design","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; CONQUEST; Software; Software engineering; Natural language processing; Programming language; History","score_opus":0.1009210075634949,"score_gpt":0.31562541546700623,"score_spread":0.21470440790351134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098969028","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1075317,0.01204417,0.6828972,0.05858428,0.0025253282,0.00016063817,0.0011760527,0.0030905537,0.13199021],"genre_scores_gemma":[0.8957162,0.0024889605,0.07865721,0.009756442,0.0006380491,0.0001460745,0.0006488255,0.0009395295,0.011008763],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.988729,0.0063518216,0.00091576733,0.0010495004,0.0024044828,0.0005493815],"domain_scores_gemma":[0.98292613,0.01031799,0.0014279499,0.0022656894,0.002648527,0.000413672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068538194,0.00095456117,0.0005859751,0.0015505989,0.0030182952,0.0039518094,0.00080715504,0.0031745217,0.00339808],"category_scores_gemma":[0.01650319,0.00084162847,0.0007841528,0.0013725001,0.011898546,0.011541732,0.0034655128,0.0051966845,0.0009182041],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008374154,0.000011585934,0.0006247544,0.00022168392,0.000013289775,0.00036412632,0.008980584,0.00037925865,0.0021213228,0.95720905,0.011223676,0.01876688],"study_design_scores_gemma":[0.000022266802,0.000093779294,0.0008875523,0.00034381493,0.000048264326,0.0016679508,0.004634989,0.0028130356,0.00613997,0.8000864,0.18319,0.000071943614],"about_ca_topic_score_codex":0.0017466231,"about_ca_topic_score_gemma":0.001592929,"teacher_disagreement_score":0.0068538194,"about_ca_system_score_codex":0.0019874773,"about_ca_system_score_gemma":0.0013612553,"threshold_uncertainty_score":0.036246896},"labels":[],"label_agreement":null},{"id":"W2098975596","doi":"10.1109/icsm.2008.4658066","title":"Supporting software evolution using adaptive change propagation heuristics","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Heuristics; Computer science; Compiler; Software; Software system; Source code; Programming language; Operating system","score_opus":0.08383529020130019,"score_gpt":0.30572281828004544,"score_spread":0.22188752807874523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2098975596","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17989634,0.00097273296,0.8019846,0.00061103527,0.00012783938,0.0006585306,0.00055972644,0.011630814,0.0035583687],"genre_scores_gemma":[0.49954504,0.00030397743,0.49739847,0.0002693865,0.000050242354,0.00032260382,0.0010751698,0.00033181589,0.00070328143],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99564314,0.0015823238,0.0005173798,0.0009273676,0.0010001247,0.00032971273],"domain_scores_gemma":[0.96830606,0.021539714,0.0029087127,0.0029498418,0.003850237,0.00044537312],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006093509,0.0013599927,0.0014408818,0.0053541036,0.0009808103,0.002245701,0.0030255658,0.0016235288,0.0010619588],"category_scores_gemma":[0.033978537,0.00079719396,0.0008915828,0.0040704375,0.00077742484,0.0029232926,0.0012472277,0.0012696366,0.0004711023],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046334515,0.00070406025,0.037268255,0.00055026513,0.0003809562,0.00029713806,0.0009542866,0.31359273,0.010911335,0.006040395,0.0068845493,0.6219527],"study_design_scores_gemma":[0.0001335597,0.00015953736,0.0038731054,0.000060521106,0.00019681455,0.00022570518,0.00024694647,0.9761705,0.008754884,0.006806609,0.0033002559,0.000071513896],"about_ca_topic_score_codex":0.010600793,"about_ca_topic_score_gemma":0.014538551,"teacher_disagreement_score":0.010600793,"about_ca_system_score_codex":0.0014709431,"about_ca_system_score_gemma":0.0028571847,"threshold_uncertainty_score":0.032225907},"labels":[],"label_agreement":null},{"id":"W2099069768","doi":"10.1109/tse.2005.106","title":"Analyzing the evolutionary history of the logical design of object-oriented software","year":2005,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software evolution; Unified Modeling Language; Object-oriented design; Software system; Class (philosophy); Programming language; Inheritance (genetic algorithm); Sequence diagram; Class diagram; Abstraction; Object-oriented programming; Sequence (biology); Software; Software engineering; Artificial intelligence; Software construction","score_opus":0.02065955642586971,"score_gpt":0.22516811738749967,"score_spread":0.20450856096162995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099069768","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7357845,0.0011515311,0.25726047,0.0005285609,0.000016967851,0.000114083596,0.00021825572,0.00014745645,0.004778227],"genre_scores_gemma":[0.8466096,0.00066548074,0.1508315,0.000060888222,0.00001217632,0.000073524265,0.00047795288,0.000053832657,0.0012150094],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99850225,0.00063524576,0.00009252697,0.00020623313,0.00050191156,0.00006183554],"domain_scores_gemma":[0.9926267,0.0035163707,0.0013162781,0.00092207664,0.0014496269,0.00016882762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032914602,0.000182372,0.00017420598,0.0030481,0.00076851377,0.0013621825,0.0004493118,0.00043592346,0.0005747259],"category_scores_gemma":[0.01522747,0.0003478626,0.00029683538,0.002226571,0.0009479161,0.0027116616,0.00059444393,0.00072121114,0.00010917307],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022457707,0.00022726982,0.3204443,0.00040068192,0.00016468343,0.00063978374,0.0082708625,0.041073862,0.020456007,0.11186485,0.00080422993,0.495429],"study_design_scores_gemma":[0.000058432106,0.0006254491,0.43315938,0.0004084705,0.00025291066,0.0017994716,0.0046213726,0.34024107,0.038779266,0.12944877,0.05042516,0.00018021766],"about_ca_topic_score_codex":0.0035250813,"about_ca_topic_score_gemma":0.0052732914,"teacher_disagreement_score":0.0035250813,"about_ca_system_score_codex":0.0016821922,"about_ca_system_score_gemma":0.0011894187,"threshold_uncertainty_score":0.01740712},"labels":[],"label_agreement":null},{"id":"W2099287333","doi":"10.1109/cnsr.2008.78","title":"Information Retrieval in Network Administration","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"National Institute for Materials Science; Research Nova Scotia; Dalhousie University","keywords":"Vocabulary; Computer science; Task (project management); Information retrieval; Quality (philosophy); Component (thermodynamics); Artificial intelligence; Natural language processing; Question answering; Term (time); Linguistics","score_opus":0.01821430077502995,"score_gpt":0.2522090357885117,"score_spread":0.23399473501348178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099287333","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03278729,0.04367867,0.80887884,0.01345341,0.0009425094,0.00084280065,0.0007486161,0.003080077,0.0955878],"genre_scores_gemma":[0.34050968,0.025931776,0.5895739,0.0023909744,0.0017645735,0.0006122328,0.0020404158,0.00062725577,0.036549207],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9916083,0.0044619814,0.00078793993,0.0010128727,0.0018449539,0.0002838401],"domain_scores_gemma":[0.9809724,0.012503031,0.0011173515,0.0025732375,0.0024142757,0.00041970707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073042368,0.0005744781,0.0011522416,0.008569747,0.0022879525,0.009721931,0.0020184424,0.002531688,0.009287552],"category_scores_gemma":[0.032322556,0.0006338403,0.0008609396,0.010267264,0.0040159635,0.017920978,0.0027486924,0.0015313821,0.0049733776],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020678523,0.00015594135,0.0031274706,0.0021156822,0.00012811297,0.0005135076,0.0034556526,0.009832369,0.006820803,0.28360784,0.044206996,0.64582884],"study_design_scores_gemma":[0.00010547243,0.00023295997,0.0048405775,0.0009939919,0.00015054006,0.0017316294,0.0030206286,0.073847346,0.012966104,0.59709066,0.30485538,0.0001646147],"about_ca_topic_score_codex":0.0050822403,"about_ca_topic_score_gemma":0.0034922413,"teacher_disagreement_score":0.009721931,"about_ca_system_score_codex":0.003298979,"about_ca_system_score_gemma":0.0020738533,"threshold_uncertainty_score":0.038628936},"labels":[],"label_agreement":null},{"id":"W2099290551","doi":"10.1145/2024587.2024595","title":"An explanatory analysis on eclipse beta-release bugs through in-process metrics","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Eclipse; Computer science; Software bug; Process (computing); Code (set theory); BETA (programming language); Software release life cycle; Software quality; Software; Software engineering; Programming language; Software development; Set (abstract data type)","score_opus":0.055828183628138565,"score_gpt":0.3148908320692412,"score_spread":0.2590626484411026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099290551","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9919431,0.00014592857,0.005896051,0.00014620286,0.0000070894034,0.000043670156,0.0009857769,0.00016862465,0.00066352304],"genre_scores_gemma":[0.9952419,0.0000673465,0.003315312,0.000014980603,0.000009995881,0.000030859508,0.0010049768,0.00003231477,0.00028254383],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945972,0.0024835467,0.0005426772,0.00055932294,0.0015811501,0.00023600025],"domain_scores_gemma":[0.7977684,0.15098312,0.027648829,0.0071385093,0.015325019,0.0011361341],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00750567,0.000715999,0.00036505086,0.0050056274,0.00039523217,0.00091023464,0.00061572506,0.00038098235,0.00094206206],"category_scores_gemma":[0.06124579,0.00027759897,0.00065217534,0.004147427,0.00040197352,0.0013172269,0.0005743073,0.00067010045,0.00017363022],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024917768,0.00014796286,0.9660854,0.00014651376,0.000172165,0.00036538442,0.0017682737,0.0019930745,0.0022011285,0.0005863463,0.0005215129,0.025763182],"study_design_scores_gemma":[0.000010678057,0.0005279931,0.98305297,0.000047280515,0.0001596602,0.0003047344,0.00076993764,0.01176052,0.0019609325,0.0002532201,0.0011258777,0.00002624049],"about_ca_topic_score_codex":0.0036120296,"about_ca_topic_score_gemma":0.0040975115,"teacher_disagreement_score":0.00750567,"about_ca_system_score_codex":0.00057762826,"about_ca_system_score_gemma":0.000790613,"threshold_uncertainty_score":0.03969425},"labels":[],"label_agreement":null},{"id":"W2099353275","doi":"10.1109/icpc.2006.19","title":"Dynamic Analysis of Software Systems using Execution Pattern Mining","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":82,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Program comprehension; Unix; Software system; Feature (linguistics); Task (project management); Software; Source code; Data mining; Software construction; Software visualization; Static program analysis; Software evolution; Software engineering; Software development; Programming language","score_opus":0.015517748902527448,"score_gpt":0.26396484215574945,"score_spread":0.248447093253222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099353275","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09269199,0.000508187,0.90030336,0.00025349803,0.000021269952,0.00025937698,0.0011629581,0.0031701662,0.0016291676],"genre_scores_gemma":[0.4304353,0.00055556314,0.5640654,0.000053738448,0.000021992864,0.00032508583,0.0032236043,0.00021475657,0.0011045542],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986607,0.000289283,0.000147419,0.000320868,0.0004923853,0.000089423775],"domain_scores_gemma":[0.9966619,0.0016942465,0.00044316423,0.00064458797,0.0004871767,0.00006897817],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010831305,0.0009050192,0.000694275,0.0050143604,0.00046394163,0.0009799323,0.0009092015,0.00053530146,0.00097504194],"category_scores_gemma":[0.0043890285,0.00033537397,0.0009259498,0.0031430542,0.00045897413,0.0016881219,0.0007170212,0.0006363677,0.00046698892],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027347356,0.00038898687,0.04065798,0.00087881077,0.00029359866,0.0009046036,0.0006208807,0.07126143,0.052362382,0.011063735,0.0019437477,0.8193504],"study_design_scores_gemma":[0.00003864515,0.0003190056,0.016852435,0.00011416323,0.00016083283,0.0012120524,0.00042163572,0.8839311,0.054389834,0.032743663,0.009756015,0.000060672995],"about_ca_topic_score_codex":0.0016942162,"about_ca_topic_score_gemma":0.0019811417,"teacher_disagreement_score":0.0050143604,"about_ca_system_score_codex":0.0004114353,"about_ca_system_score_gemma":0.0008410167,"threshold_uncertainty_score":0.005728245},"labels":[],"label_agreement":null},{"id":"W2099499926","doi":"10.1109/ccece.1993.332408","title":"Evaluating expert systems by formal metrics","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Metric (unit); Measure (data warehouse); Computer science; Expert system; Artificial intelligence; Software engineering; Machine learning; Data mining; Engineering","score_opus":0.10367739302401864,"score_gpt":0.3442505361667756,"score_spread":0.24057314314275696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099499926","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10232598,0.0012407538,0.88643986,0.0007239921,0.00009812367,0.0005257191,0.00039925202,0.0007401224,0.0075062574],"genre_scores_gemma":[0.5288203,0.0008881568,0.46702495,0.00011078697,0.00012685034,0.0008604979,0.0010355302,0.0001805927,0.0009524275],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9547998,0.020742653,0.0044385204,0.0019664825,0.017172368,0.0008801174],"domain_scores_gemma":[0.8505357,0.102786124,0.015526281,0.0110218,0.018487528,0.0016424807],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018366547,0.0015372473,0.0009620406,0.0075610867,0.0006215874,0.0053146877,0.0012246416,0.0015110709,0.0019815888],"category_scores_gemma":[0.13121134,0.0004319571,0.0007712186,0.0045529916,0.002688282,0.0085754655,0.0022549864,0.0012032578,0.0002859737],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017217842,0.00026579935,0.017514331,0.00084049575,0.00026270753,0.00015276088,0.0009389571,0.47632462,0.005104216,0.24090518,0.0025917254,0.25492713],"study_design_scores_gemma":[0.000060314727,0.00046739297,0.0046485234,0.00023456854,0.000076222495,0.00016751936,0.0003433236,0.8280372,0.0049822237,0.15388219,0.007021266,0.00007923044],"about_ca_topic_score_codex":0.0019908717,"about_ca_topic_score_gemma":0.0022280696,"teacher_disagreement_score":0.018366547,"about_ca_system_score_codex":0.003142255,"about_ca_system_score_gemma":0.002348463,"threshold_uncertainty_score":0.09713274},"labels":[],"label_agreement":null},{"id":"W2099526349","doi":"10.5555/2664398.2664418","title":"We have all of the clones, now what?: toward integrating clone analysis into software quality assessment","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"clone (Java method); Consistency (knowledge bases); Cloning (programming); Software; Software maintenance; Computer science; Software development; Software engineering; Software quality; Software evolution; Quality (philosophy); Data science; Software construction; Artificial intelligence; Biology; Programming language","score_opus":0.059046826759216575,"score_gpt":0.3625117195191857,"score_spread":0.30346489275996913,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099526349","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07681553,0.0014474667,0.9078612,0.006539009,0.00012682306,0.000236716,0.00007306572,0.0019205519,0.004979633],"genre_scores_gemma":[0.2429552,0.0007636245,0.7535756,0.00079129456,0.00007240695,0.00015257554,0.00012636007,0.0002028667,0.0013599875],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9883782,0.0055566053,0.0009312315,0.0013940566,0.003346243,0.00039370923],"domain_scores_gemma":[0.94295937,0.022679513,0.008031777,0.005934009,0.018917859,0.001477524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015885916,0.0008202325,0.0009022128,0.007413723,0.0019015771,0.009597805,0.0019705587,0.0024492922,0.0014454225],"category_scores_gemma":[0.06707335,0.0008120133,0.00067442056,0.004795424,0.005309621,0.011991528,0.004213045,0.002904783,0.00078679575],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020802021,0.00040916965,0.101024464,0.00075716357,0.00012864808,0.0003972561,0.010204688,0.008363507,0.016941767,0.11068377,0.006684654,0.74419683],"study_design_scores_gemma":[0.00009264394,0.00093627675,0.06144344,0.0025912495,0.00040251444,0.003135557,0.019851057,0.39514384,0.045886446,0.41255394,0.057453852,0.00050912495],"about_ca_topic_score_codex":0.0048140204,"about_ca_topic_score_gemma":0.0038640795,"teacher_disagreement_score":0.015885916,"about_ca_system_score_codex":0.0019893793,"about_ca_system_score_gemma":0.003606162,"threshold_uncertainty_score":0.08401376},"labels":[],"label_agreement":null},{"id":"W2099699509","doi":"10.1109/wcre.2003.1287257","title":"Leveraging visio for adoption-centric reverse engineering tools","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Reverse engineering; Computer science; Software engineering; Leverage (statistics); Visualization; Software; Domain (mathematical analysis); Programming language; Data mining; Artificial intelligence","score_opus":0.02896057617612029,"score_gpt":0.2571318590002598,"score_spread":0.2281712828241395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099699509","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04128316,0.00030719128,0.91105115,0.0015926475,0.000082848164,0.00021393463,0.00015529337,0.028818076,0.016495634],"genre_scores_gemma":[0.177133,0.00042703818,0.81412035,0.00026360241,0.000057609428,0.00023912446,0.00040849898,0.002843493,0.0045073694],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961526,0.0013267524,0.00038914598,0.0006439563,0.0013103663,0.00017708683],"domain_scores_gemma":[0.9654267,0.018667206,0.002904514,0.008690703,0.0034932084,0.0008176666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072856625,0.0012318774,0.0005454432,0.0038877272,0.00055663456,0.004783975,0.0015967762,0.0012292238,0.0027476088],"category_scores_gemma":[0.035559583,0.0009071105,0.00076813926,0.0016566245,0.0012192973,0.009296832,0.0044747437,0.0024245288,0.0013123181],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033721444,0.0004549329,0.011953142,0.0010899284,0.00021317987,0.0017648228,0.010195443,0.009247609,0.068952635,0.060673155,0.0111965,0.8239214],"study_design_scores_gemma":[0.00032144942,0.001336222,0.01685935,0.0019037776,0.0005245764,0.0061141066,0.005060359,0.20114286,0.13019809,0.14096026,0.49500662,0.00057243655],"about_ca_topic_score_codex":0.0008935623,"about_ca_topic_score_gemma":0.0019094169,"teacher_disagreement_score":0.0072856625,"about_ca_system_score_codex":0.0007238973,"about_ca_system_score_gemma":0.0014869254,"threshold_uncertainty_score":0.038530767},"labels":[],"label_agreement":null},{"id":"W2099706779","doi":"10.1109/icsm.2009.5306335","title":"Searching and skimming: An exploratory study","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Formative assessment; Task (project management); Source code; Program comprehension; Exploratory research; Code (set theory); Task analysis; Software; Code review; Human–computer interaction; Software engineering; Static program analysis; Software development; World Wide Web; Data science; Software system; Programming language; Set (abstract data type); Engineering","score_opus":0.034353527006351046,"score_gpt":0.31024594100015396,"score_spread":0.27589241399380293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099706779","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9961964,0.00011677399,0.001574392,0.0002573735,0.000009222816,0.0004105375,0.00010867812,0.00004802419,0.0012785619],"genre_scores_gemma":[0.9901738,0.00045713567,0.0055266493,0.0006412963,0.000038814305,0.0008919479,0.00021685415,0.00007091517,0.0019825865],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99457693,0.002998033,0.000445634,0.0005523042,0.0008105636,0.0006164835],"domain_scores_gemma":[0.9525315,0.03891751,0.0025220225,0.0015144586,0.0027647521,0.0017497048],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.009328264,0.0010059663,0.0012637494,0.0025207885,0.005813504,0.0033432166,0.0019984932,0.003306221,0.0019261002],"category_scores_gemma":[0.049520284,0.001323613,0.00050710386,0.0014642118,0.0032852886,0.0046383925,0.0030479217,0.003050141,0.0007462196],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037565146,0.0023783487,0.026832454,0.0006169529,0.000027527118,0.0032585023,0.94038355,0.00011427597,0.008860105,0.00062651193,0.00094497687,0.01558115],"study_design_scores_gemma":[0.0001881045,0.00442298,0.06340465,0.00054635,0.00006966773,0.005129472,0.8991248,0.0017719984,0.0058794464,0.0016117074,0.01762351,0.0002272629],"about_ca_topic_score_codex":0.0028053296,"about_ca_topic_score_gemma":0.004306203,"teacher_disagreement_score":0.9966568,"about_ca_system_score_codex":0.0009388465,"about_ca_system_score_gemma":0.0019196321,"threshold_uncertainty_score":0.049333155},"labels":[],"label_agreement":null},{"id":"W2099784938","doi":"","title":"Budgeting for Information Security and ROI Approach","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Risk analysis (engineering); Metric (unit); Cost–benefit analysis; Business; Investment (military); Information security; Computer science; Actuarial science; Computer security; Marketing","score_opus":0.00984437173130332,"score_gpt":0.23914312567515308,"score_spread":0.22929875394384977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099784938","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040068045,0.0026075249,0.8360625,0.0039151125,0.00019635863,0.0010390016,0.00092978956,0.0003511076,0.11483056],"genre_scores_gemma":[0.77776724,0.0020328441,0.20358464,0.00020032236,0.000121946374,0.0012487476,0.00049702066,0.0001137994,0.014433382],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9919194,0.0053484086,0.00028124388,0.00035468067,0.0016365724,0.00045982114],"domain_scores_gemma":[0.99531364,0.002649627,0.00049258035,0.0002526532,0.0011092776,0.00018215325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069204364,0.00133539,0.0007922273,0.0034596962,0.00060318253,0.0034784374,0.001272172,0.0010034632,0.008459998],"category_scores_gemma":[0.016327687,0.0006049618,0.00060473476,0.0034401156,0.00090397336,0.003126997,0.0012515552,0.0012025412,0.0007007278],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012523039,0.00010543626,0.0038980816,0.000515035,0.00014053215,0.00024257475,0.00043572733,0.2726142,0.0011288538,0.5171226,0.009171008,0.19450077],"study_design_scores_gemma":[0.000058654416,0.00025616767,0.006248571,0.0007413274,0.00011919989,0.0004177561,0.0009184567,0.6126793,0.00214134,0.31789944,0.05842815,0.00009157869],"about_ca_topic_score_codex":0.0038162374,"about_ca_topic_score_gemma":0.0037844756,"teacher_disagreement_score":0.008459998,"about_ca_system_score_codex":0.004899404,"about_ca_system_score_gemma":0.0038747375,"threshold_uncertainty_score":0.03659922},"labels":[],"label_agreement":null},{"id":"W2099901245","doi":"10.1109/wcre.2008.15","title":"Retrieving Task-Related Clusters from Change History","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Task (project management); Software maintenance; Software; Source code; Software system; Software engineering; Data science; Programming language; Engineering; Systems engineering","score_opus":0.05401685614114267,"score_gpt":0.23102986197661382,"score_spread":0.17701300583547114,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099901245","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8325198,0.0028708184,0.12866598,0.0008587259,0.00016177107,0.0010147254,0.019202562,0.007123703,0.0075818775],"genre_scores_gemma":[0.8313597,0.0010198642,0.13267218,0.0001145292,0.000096924006,0.00044214522,0.030617857,0.0008370976,0.0028397604],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981364,0.00024026734,0.000230845,0.00063590205,0.0005760675,0.00018057234],"domain_scores_gemma":[0.97630465,0.010021312,0.0034968175,0.0038963454,0.0051333243,0.0011474724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019555627,0.0014071864,0.0013995549,0.020962406,0.001271582,0.002717101,0.0014557295,0.0012570119,0.002118639],"category_scores_gemma":[0.029823465,0.000824699,0.0009628458,0.014126418,0.0005535173,0.0047169137,0.002098116,0.0012537774,0.0016469589],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001007343,0.00055038277,0.3590477,0.0010890863,0.00030165305,0.0009780775,0.008422933,0.013372363,0.014336166,0.0030168083,0.015921935,0.58195555],"study_design_scores_gemma":[0.00017055859,0.00090398773,0.6897707,0.00058017217,0.00073532807,0.001829024,0.009664152,0.19683003,0.024210906,0.0242222,0.050546672,0.0005362513],"about_ca_topic_score_codex":0.029696885,"about_ca_topic_score_gemma":0.03722609,"teacher_disagreement_score":0.029696885,"about_ca_system_score_codex":0.0013898534,"about_ca_system_score_gemma":0.0024462522,"threshold_uncertainty_score":0.059048057},"labels":[],"label_agreement":null},{"id":"W2099939057","doi":"10.2139/ssrn.1447584","title":"Optimal Enhancement and Lifetime of Software Systems: A Control Theoretic Analysis","year":2010,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software; Control (management); Econometrics; Mathematics; Artificial intelligence","score_opus":0.003115216860530052,"score_gpt":0.22439872637368685,"score_spread":0.2212835095131568,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099939057","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34562072,0.002009454,0.6329366,0.0019296322,0.00007352141,0.00008688208,0.0001831134,0.0002212945,0.016938811],"genre_scores_gemma":[0.9831404,0.00072214723,0.011262061,0.00007723338,0.00006880628,0.000058523023,0.00004125474,0.000056574165,0.004572978],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918526,0.00029964422,0.000031059848,0.00013709346,0.00015256762,0.00019430355],"domain_scores_gemma":[0.9888782,0.008282701,0.0010298057,0.00039151317,0.0008972408,0.00052052294],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003022712,0.00084961916,0.001095761,0.0014402015,0.00059092214,0.0017259119,0.0015148426,0.0012090587,0.003572055],"category_scores_gemma":[0.01618827,0.00064803095,0.0007147119,0.0007282475,0.0026570861,0.0045923134,0.0012546524,0.0012866068,0.00015308885],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021377968,0.00013186093,0.0013467167,0.00021015898,0.00005806135,0.00014575855,0.0003524228,0.43204668,0.0067876386,0.53395873,0.0014906427,0.023257516],"study_design_scores_gemma":[0.00002181666,0.00009634871,0.00080070645,0.000027515689,0.000032766366,0.000085677646,0.00010257,0.793505,0.00096084893,0.20383061,0.0005150105,0.000021089283],"about_ca_topic_score_codex":0.0019358333,"about_ca_topic_score_gemma":0.0011078311,"teacher_disagreement_score":0.003572055,"about_ca_system_score_codex":0.0020412877,"about_ca_system_score_gemma":0.0011879966,"threshold_uncertainty_score":0.015985847},"labels":[],"label_agreement":null},{"id":"W2100060170","doi":"10.1109/icstw.2009.18","title":"A Mutation/Injection-Based Automatic Framework for Evaluating Code Clone Detection Tools","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":195,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Computer science; Benchmark (surveying); Precision and recall; Software; Software maintenance; Code (set theory); Data mining; Mutation; Machine learning; Software system; Software engineering; Artificial intelligence; Programming language; Set (abstract data type); Biology","score_opus":0.052369893842208914,"score_gpt":0.360897314909121,"score_spread":0.3085274210669121,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100060170","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09782159,0.00055569224,0.8786755,0.00013626825,0.000036015328,0.0015737271,0.00074972067,0.018287506,0.002164036],"genre_scores_gemma":[0.27433515,0.00009363458,0.7227397,0.000058458718,0.000020954178,0.0011195376,0.00090376486,0.00029310185,0.00043564042],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97093844,0.007983164,0.003164149,0.0030251504,0.014073932,0.00081527734],"domain_scores_gemma":[0.92684686,0.03199418,0.013774001,0.008965603,0.017552987,0.00086641073],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.019943383,0.0026162907,0.002855642,0.017511863,0.0015968455,0.0030823832,0.0037389372,0.0027140626,0.0011954927],"category_scores_gemma":[0.07022418,0.0009755414,0.0016594913,0.005556416,0.0022122178,0.0039822757,0.0023112865,0.0015597912,0.0005798037],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008919647,0.002069632,0.07359952,0.0011665708,0.00072197645,0.0003585666,0.0010213424,0.12628263,0.09522365,0.015277234,0.0045637614,0.6788232],"study_design_scores_gemma":[0.00023362138,0.002374292,0.036445484,0.00014082555,0.0002273231,0.0007895927,0.00020451071,0.8889894,0.060527127,0.005882257,0.0038558245,0.00032977285],"about_ca_topic_score_codex":0.0073654484,"about_ca_topic_score_gemma":0.006592975,"teacher_disagreement_score":0.98005664,"about_ca_system_score_codex":0.0026935951,"about_ca_system_score_gemma":0.0042642965,"threshold_uncertainty_score":0.10547197},"labels":[],"label_agreement":null},{"id":"W2100136054","doi":"10.1109/wcre.2002.1173078","title":"An extensible tool for source code representation using XML","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; ASCII; Programming language; XML; Plain text; XML framework; Source code; Streaming XML; Document Structure Description; Static program analysis; Software; Software engineering; Database; Software development; World Wide Web; Operating system; Encryption","score_opus":0.05904711780685137,"score_gpt":0.3484391893904717,"score_spread":0.28939207158362035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100136054","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013563875,0.00020140829,0.8615996,0.00013410927,0.00012264217,0.00037197553,0.00319679,0.12982136,0.00319572],"genre_scores_gemma":[0.0179199,0.000924108,0.9250304,0.00028618998,0.00009741703,0.001188245,0.023297073,0.020758918,0.010497592],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977671,0.0003225854,0.0005832216,0.00032083658,0.0009061668,0.00010023973],"domain_scores_gemma":[0.9961133,0.0015636864,0.00031338484,0.001060416,0.0007694507,0.00017969937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040411684,0.0019359933,0.00084658957,0.0053884787,0.00089836423,0.0036245964,0.0043994277,0.0018285181,0.017052844],"category_scores_gemma":[0.010230122,0.0019697216,0.001877503,0.004057175,0.0006533587,0.004628404,0.0033460204,0.0034429876,0.009946096],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005779306,0.00037805949,0.0020570909,0.0014977811,0.00019257456,0.00236966,0.0015497658,0.010909036,0.034701142,0.059573326,0.10899404,0.77719957],"study_design_scores_gemma":[0.00042166424,0.00030735438,0.0022179314,0.0008984529,0.00017563622,0.0032855205,0.0002829047,0.100445785,0.06425759,0.04607646,0.7812721,0.00035859869],"about_ca_topic_score_codex":0.0016960074,"about_ca_topic_score_gemma":0.0012624695,"teacher_disagreement_score":0.017052844,"about_ca_system_score_codex":0.00057746563,"about_ca_system_score_gemma":0.0013189907,"threshold_uncertainty_score":0.057047486},"labels":[],"label_agreement":null},{"id":"W2100348239","doi":"10.1109/icpc.2011.42","title":"Trust-Based Requirements Traceability","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Traceability; Requirements traceability; Computer science; Source code; Documentation; Precision and recall; Set (abstract data type); Code (set theory); Data mining; Software engineering; Information retrieval; Requirements analysis; Programming language; Software; Requirement","score_opus":0.07953031252904559,"score_gpt":0.29006904304557946,"score_spread":0.21053873051653388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100348239","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03671804,0.00023491502,0.94728154,0.00062069006,0.000035728837,0.00040414964,0.00054136285,0.0066855573,0.0074780285],"genre_scores_gemma":[0.62603354,0.00022492114,0.3659887,0.00020307695,0.000028103237,0.00037051077,0.0021559391,0.0008194487,0.0041756723],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9728223,0.0093143955,0.0020258198,0.0036591373,0.011269566,0.0009087156],"domain_scores_gemma":[0.92863923,0.028896417,0.0072385003,0.020734083,0.013759985,0.0007317926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012167974,0.0015635571,0.00095841964,0.0060481722,0.0011916346,0.0038953093,0.0025232155,0.001410382,0.004383069],"category_scores_gemma":[0.09131779,0.0008550093,0.0017100989,0.0026569283,0.0014174599,0.007039822,0.0038006469,0.0025862837,0.0016550936],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004907781,0.00058989605,0.014886481,0.0008691318,0.00028196565,0.00059830846,0.0031203683,0.14842072,0.018286886,0.051285066,0.00783751,0.7533329],"study_design_scores_gemma":[0.00005324151,0.00029541223,0.004865364,0.00019267862,0.000111773516,0.0004256574,0.00082021987,0.8766406,0.036064472,0.06282582,0.01757458,0.0001301221],"about_ca_topic_score_codex":0.0115621975,"about_ca_topic_score_gemma":0.008376036,"teacher_disagreement_score":0.012167974,"about_ca_system_score_codex":0.0031493604,"about_ca_system_score_gemma":0.0036327715,"threshold_uncertainty_score":0.0643512},"labels":[],"label_agreement":null},{"id":"W2100352505","doi":"10.1109/tai.2003.1250170","title":"Analysis of software maintenance data using multi-technique approach","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Data mining; Software; Data modeling; Process (computing); Bayesian network; Software engineering; Software development; Software maintenance; Decision tree; Set (abstract data type); Software development process; Data science; Machine learning","score_opus":0.08555855559529078,"score_gpt":0.3258004628075996,"score_spread":0.24024190721230881,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100352505","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08306129,0.0007309789,0.91242594,0.00025749468,0.000032868353,0.00037496496,0.0007154309,0.0008174273,0.0015836981],"genre_scores_gemma":[0.35557672,0.0005426594,0.6411607,0.000064998414,0.000040040082,0.0004884084,0.0014544495,0.00009886031,0.0005730767],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9920306,0.0033789785,0.0007090172,0.00074816745,0.0028007268,0.00033255058],"domain_scores_gemma":[0.96841305,0.023242086,0.0018919213,0.0020458733,0.0041441103,0.00026299548],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060069626,0.0008009234,0.0014011861,0.014309389,0.0006909852,0.001804254,0.00096306915,0.0012484774,0.0015371997],"category_scores_gemma":[0.021510048,0.00041891588,0.002313335,0.0075818575,0.00045563176,0.0022933881,0.0015884413,0.0014130824,0.0005322504],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088882534,0.00077592925,0.06192939,0.0015716277,0.0015616457,0.0015757495,0.003361663,0.08242024,0.0377176,0.013177675,0.0019860321,0.7930338],"study_design_scores_gemma":[0.00009654087,0.0013380365,0.06155876,0.00031416747,0.00059484283,0.0019376962,0.0022986233,0.8721751,0.02125975,0.027721075,0.010410225,0.00029524832],"about_ca_topic_score_codex":0.002485663,"about_ca_topic_score_gemma":0.002699534,"teacher_disagreement_score":0.014309389,"about_ca_system_score_codex":0.0007643496,"about_ca_system_score_gemma":0.0009863587,"threshold_uncertainty_score":0.031768203},"labels":[],"label_agreement":null},{"id":"W2100510954","doi":"10.1049/ic:20040215","title":"Program navigation analysis to support task-aware software development environments","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Debugging; Task (project management); Software engineering; Software; Visualization; Software development; Human–computer interaction; Program comprehension; Interface (matter); User interface; Task analysis; Application programming interface; Program analysis; Software system; Programming language; Systems engineering; Artificial intelligence; Operating system; Engineering","score_opus":0.014947218303547926,"score_gpt":0.2785942854294169,"score_spread":0.26364706712586894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100510954","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022289066,0.00021856268,0.9613656,0.0002435159,0.000025569763,0.00014985134,0.00016395209,0.014387809,0.0011559695],"genre_scores_gemma":[0.13649146,0.00020499452,0.860558,0.00009389407,0.000025732152,0.00020217277,0.0004985388,0.001003435,0.00092175],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99802774,0.0007008022,0.00017814584,0.00034305808,0.00063937873,0.0001109488],"domain_scores_gemma":[0.9876416,0.0067610517,0.001595761,0.0018521607,0.0018795542,0.00026989196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025487682,0.00081079017,0.0005699397,0.0024398072,0.00081841456,0.0014575211,0.0013496137,0.00065035734,0.001323422],"category_scores_gemma":[0.015112605,0.0005733715,0.00063606055,0.0017646597,0.00080778333,0.002650444,0.0012593217,0.001413406,0.00051121286],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055476325,0.00033266435,0.018500876,0.0008285039,0.0001561155,0.0005424416,0.0045765396,0.019422276,0.042100366,0.026057296,0.010756674,0.8761715],"study_design_scores_gemma":[0.00017417352,0.0005456074,0.01697803,0.00058712304,0.0003488718,0.0016854139,0.0013868662,0.67829365,0.09511003,0.103171736,0.10142241,0.00029617376],"about_ca_topic_score_codex":0.003314246,"about_ca_topic_score_gemma":0.0051917937,"teacher_disagreement_score":0.003314246,"about_ca_system_score_codex":0.00055098074,"about_ca_system_score_gemma":0.0022719966,"threshold_uncertainty_score":0.013479352},"labels":[],"label_agreement":null},{"id":"W2100568614","doi":"10.1002/spe.1001","title":"The use of search‐based optimization techniques to schedule and staff software projects: an approach and an empirical study","year":2011,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Staffing; Computer science; Schedule; Software project management; Software; Queueing theory; Task (project management); Scheduling (production processes); Fragmentation (computing); Operations research; Project management; Project planning; Operations management; Software development; Systems engineering; Software construction; Engineering","score_opus":0.12405668412549276,"score_gpt":0.3731303624344601,"score_spread":0.24907367830896737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100568614","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89625317,0.00059046625,0.09851631,0.00024844284,0.000010767941,0.0002848449,0.00009297275,0.00018173883,0.0038212093],"genre_scores_gemma":[0.940183,0.0002395304,0.058673177,0.000018389013,0.0000051489105,0.00023552352,0.00010950104,0.000024049277,0.0005116306],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9965938,0.0026340915,0.00013942191,0.00013659883,0.00042278165,0.0000732262],"domain_scores_gemma":[0.9549844,0.041761257,0.0011563374,0.000827001,0.0010809967,0.00019000426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005622545,0.00053032703,0.0004666467,0.0015028829,0.00039946224,0.0006856593,0.00080076244,0.0006847519,0.0016230257],"category_scores_gemma":[0.029028542,0.00036264208,0.00039027596,0.00205936,0.00083394407,0.0013262989,0.00050227134,0.00066663494,0.0002140022],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012516316,0.0042098234,0.04240977,0.00088905025,0.00022237895,0.00020595013,0.002654239,0.638743,0.0040191263,0.012944435,0.0020906725,0.2903599],"study_design_scores_gemma":[0.00017921087,0.0012930486,0.013754257,0.00005112524,0.000056526995,0.00013672003,0.0007474102,0.97844625,0.0018537282,0.002298921,0.0011494424,0.00003336197],"about_ca_topic_score_codex":0.005895086,"about_ca_topic_score_gemma":0.00409968,"teacher_disagreement_score":0.005895086,"about_ca_system_score_codex":0.0009450494,"about_ca_system_score_gemma":0.0012428404,"threshold_uncertainty_score":0.029735267},"labels":[],"label_agreement":null},{"id":"W2100765983","doi":"10.1109/qsic.2007.4385497","title":"Automatic Quality Assessment of SRS Text by Means of a Decision-Tree-Based Text Classifier","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Classifier (UML); Ambiguity; Decision tree; Software quality; Decision tree learning; Software; Quality (philosophy); Natural language; Artificial intelligence; Software requirements; Software requirements specification; Software engineering; Natural language processing; Information retrieval; Data mining; Software development; Software construction; Programming language","score_opus":0.037324698259074716,"score_gpt":0.36230089927896036,"score_spread":0.32497620101988567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100765983","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12743223,0.00043414786,0.85717505,0.00047619606,0.00012636681,0.00042670898,0.0015288499,0.010634037,0.0017664732],"genre_scores_gemma":[0.39212224,0.00022140995,0.6012074,0.00010224295,0.000106880674,0.0003007338,0.004327875,0.00023389967,0.0013772885],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967541,0.00079712516,0.000552964,0.0005886109,0.0011488185,0.00015841279],"domain_scores_gemma":[0.98386127,0.008314347,0.0015820436,0.0007662135,0.0051201833,0.00035585425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030186323,0.00075357076,0.0012724191,0.0055194255,0.00067125703,0.0017464793,0.0012566395,0.0010992215,0.0013151435],"category_scores_gemma":[0.015350529,0.00021066019,0.00073895295,0.0023936601,0.00041198294,0.0021354398,0.0005563611,0.0008773174,0.0016376807],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006578096,0.00033851466,0.014330073,0.0005648221,0.00009496803,0.0006057322,0.0007901405,0.013436099,0.09335123,0.0027970867,0.00867094,0.86436254],"study_design_scores_gemma":[0.00007394005,0.00029661696,0.011896352,0.00009627401,0.00015806661,0.0005287426,0.0003558743,0.9191168,0.05692111,0.0044789594,0.006004389,0.00007281594],"about_ca_topic_score_codex":0.002384263,"about_ca_topic_score_gemma":0.0019602524,"teacher_disagreement_score":0.0055194255,"about_ca_system_score_codex":0.0006487843,"about_ca_system_score_gemma":0.0010374903,"threshold_uncertainty_score":0.01596427},"labels":[],"label_agreement":null},{"id":"W2100925270","doi":"10.1007/s10664-011-9171-y","title":"An exploratory study of the impact of antipatterns on class change- and fault-proneness","year":2011,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":394,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Odds; Computer science; Machine learning; Logistic regression","score_opus":0.10704019854817991,"score_gpt":0.3282979284401009,"score_spread":0.221257729891921,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100925270","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9993175,0.0000125178385,0.00021764157,0.000025158643,0.0000010560602,0.000013139981,0.00008861933,0.000006691472,0.00031764508],"genre_scores_gemma":[0.99922657,0.000009990315,0.00038082548,0.000013955653,0.000002872399,0.000023805338,0.0000962269,0.00000447281,0.00024136226],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99599624,0.0025647138,0.00018712578,0.0004160717,0.000590083,0.00024577827],"domain_scores_gemma":[0.69253397,0.2725572,0.020063499,0.0073995464,0.0039538457,0.0034918315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051322803,0.00033518646,0.00029389438,0.0008750966,0.0005571765,0.0007952459,0.001108225,0.00074310496,0.004539121],"category_scores_gemma":[0.0632035,0.00027673985,0.00050901674,0.0010908948,0.0009453421,0.0015443212,0.0008688805,0.0014162591,0.00041090674],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003505899,0.007942014,0.95250803,0.00012446489,0.00033683973,0.00038366194,0.0028377492,0.0018293514,0.0054524704,0.00094664283,0.0003715011,0.023761254],"study_design_scores_gemma":[0.00013548908,0.0042818324,0.9864547,0.000010874791,0.0001302777,0.00019047131,0.0016180152,0.0045029926,0.0017465068,0.0004994943,0.0004092376,0.0000200329],"about_ca_topic_score_codex":0.0022939949,"about_ca_topic_score_gemma":0.0034595632,"teacher_disagreement_score":0.0051322803,"about_ca_system_score_codex":0.00057457964,"about_ca_system_score_gemma":0.00080931,"threshold_uncertainty_score":0.027142406},"labels":[],"label_agreement":null},{"id":"W2101121486","doi":"10.1109/compsac.2011.69","title":"Reasoning about Global Clones: Scalable Semantic Clone Detection","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; Concordia University","funders":"","keywords":"Computer science; Semantic reasoner; clone (Java method); Context (archaeology); Scalability; Semantic Web; SPARQL; Data mining; Software engineering; Information retrieval; World Wide Web; Database; Artificial intelligence; RDF","score_opus":0.020031910927288482,"score_gpt":0.24921180941698684,"score_spread":0.22917989848969836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101121486","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06367117,0.00030472563,0.91914517,0.00034566113,0.000030645777,0.00016430252,0.0004668041,0.014870695,0.0010008115],"genre_scores_gemma":[0.32653022,0.00017982483,0.6691797,0.00018283197,0.00002983893,0.00012983238,0.0017407549,0.0006700784,0.0013569166],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963426,0.00054369244,0.00028855208,0.0008811129,0.0017258676,0.00021816602],"domain_scores_gemma":[0.9898201,0.00507282,0.001108003,0.0021910025,0.0015734297,0.0002346252],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028936588,0.0011961279,0.0014338073,0.004118919,0.0011428174,0.0025248001,0.0023759536,0.0018114514,0.0012631882],"category_scores_gemma":[0.014079584,0.0005787478,0.0017182889,0.0031076237,0.0011072835,0.005021412,0.002901182,0.0013615025,0.00051122846],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043147648,0.0003034922,0.027069015,0.00044978104,0.00027217838,0.00082854583,0.0013482119,0.077687636,0.028361717,0.014821071,0.008434063,0.8399929],"study_design_scores_gemma":[0.000083655745,0.00010758413,0.0033235964,0.000042117856,0.00016386638,0.0006129103,0.00048494257,0.9195247,0.0342126,0.035098065,0.006284445,0.00006145382],"about_ca_topic_score_codex":0.006322196,"about_ca_topic_score_gemma":0.00676586,"teacher_disagreement_score":0.006322196,"about_ca_system_score_codex":0.0010748821,"about_ca_system_score_gemma":0.0019377172,"threshold_uncertainty_score":0.015303314},"labels":[],"label_agreement":null},{"id":"W2101186143","doi":"10.1109/compsac.2008.173","title":"Quantifying Security in Secure Software Development Phases","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Systems development life cycle; Software security assurance; Computer science; Security bug; Secure coding; Vulnerability (computing); Software development; Artifact (error); Software development process; Computer security; Software; Software engineering; Information security; Security service; Operating system","score_opus":0.05958949005698115,"score_gpt":0.2995886240484866,"score_spread":0.2399991339915054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101186143","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7469606,0.0010137771,0.23956566,0.00018214961,0.00002599009,0.00026311542,0.0003852349,0.0004916863,0.011111829],"genre_scores_gemma":[0.9328434,0.00024852107,0.06578544,0.0000166027,0.000009073312,0.0001295332,0.0002570399,0.00005505702,0.00065532833],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9908444,0.0019752318,0.0009852176,0.0007204597,0.00498038,0.0004943051],"domain_scores_gemma":[0.96400374,0.015649194,0.010864916,0.0039855926,0.004753971,0.0007426214],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004803111,0.0009040888,0.0005181449,0.0061955494,0.00078696257,0.0018632838,0.0006173807,0.00078477064,0.00088974426],"category_scores_gemma":[0.029758641,0.0004527753,0.00070188876,0.0028054789,0.0017263292,0.0048607816,0.0027634113,0.0007224731,0.00025366907],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085616624,0.00043796303,0.32330337,0.00081585854,0.00044767995,0.0008735431,0.0035868436,0.20134033,0.0478855,0.0755812,0.001121273,0.3437502],"study_design_scores_gemma":[0.00005658447,0.001695693,0.19211976,0.00040395698,0.00042305372,0.002141543,0.0025342563,0.5073976,0.11889205,0.16303596,0.011018094,0.00028145086],"about_ca_topic_score_codex":0.0012012419,"about_ca_topic_score_gemma":0.0013007493,"teacher_disagreement_score":0.0061955494,"about_ca_system_score_codex":0.0012567617,"about_ca_system_score_gemma":0.0011044814,"threshold_uncertainty_score":0.025401533},"labels":[],"label_agreement":null},{"id":"W2101251407","doi":"10.1109/wcre.2006.1","title":"\"Cloning Considered Harmful\" Considered Harmful","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":251,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cloning (programming); Computer science; Code (set theory); Testbed; Quality (philosophy); Connotation; Software; Programming language; Risk analysis (engineering); World Wide Web","score_opus":0.018017805119319278,"score_gpt":0.24723735606038752,"score_spread":0.22921955094106825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101251407","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34014285,0.011467104,0.34522158,0.084679045,0.00439023,0.00062275765,0.00021328602,0.0025001185,0.21076308],"genre_scores_gemma":[0.8809345,0.0042437413,0.063284166,0.018350609,0.0011220055,0.0002834401,0.0001216152,0.00090979174,0.030750122],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9316385,0.034126766,0.004127017,0.0049836775,0.02278255,0.0023415834],"domain_scores_gemma":[0.8594524,0.06692585,0.021608574,0.027239833,0.022186073,0.0025872476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023972576,0.0010439064,0.0009087655,0.002950588,0.0063443063,0.007942971,0.0023909484,0.005784452,0.0044296966],"category_scores_gemma":[0.091348335,0.0007585215,0.0010777664,0.0040566665,0.019138113,0.012109008,0.0062137647,0.0050329603,0.0012893424],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003136399,0.000116637915,0.034917362,0.0015346564,0.00016159416,0.0062542567,0.061925445,0.002426562,0.018218618,0.6574765,0.021139974,0.19551471],"study_design_scores_gemma":[0.000057323683,0.0005704966,0.01689723,0.0026314345,0.00056070864,0.020499568,0.040808257,0.007072147,0.031007249,0.323338,0.5562195,0.00033806043],"about_ca_topic_score_codex":0.0024700945,"about_ca_topic_score_gemma":0.0027475113,"teacher_disagreement_score":0.023972576,"about_ca_system_score_codex":0.005596736,"about_ca_system_score_gemma":0.0051901164,"threshold_uncertainty_score":0.12678057},"labels":[],"label_agreement":null},{"id":"W2101341488","doi":"10.1109/msr.2009.5069477","title":"MapReduce as a general framework to support research in Mining Software Repositories (MSR)","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Debugging; Software deployment; Eclipse; Reuse; Software; Process (computing); Software engineering; Source code; Operating system; Resource (disambiguation); Database; Distributed computing","score_opus":0.05066098468250915,"score_gpt":0.3755869231365201,"score_spread":0.32492593845401097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101341488","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070470036,0.00050310796,0.97022444,0.0014280529,0.00014190486,0.0003558205,0.000490052,0.016002001,0.0038076416],"genre_scores_gemma":[0.087670095,0.0008135894,0.9051417,0.00039129084,0.00012584012,0.00043355153,0.0014710509,0.000879008,0.0030739189],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9976526,0.0007097949,0.00017977328,0.000350573,0.0009341848,0.00017303604],"domain_scores_gemma":[0.9969453,0.0007312106,0.00018384376,0.0012134846,0.0005468396,0.00037931596],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055176667,0.0005366594,0.00073819427,0.0017467154,0.0010688087,0.002669693,0.0032184066,0.0008327588,0.0012443637],"category_scores_gemma":[0.005472772,0.000692989,0.0012036011,0.0019816076,0.0010510009,0.0032742568,0.002865165,0.0014400115,0.0009214104],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054265704,0.0005013235,0.006799535,0.001658213,0.00041796805,0.0010294007,0.002163325,0.07062575,0.029612657,0.18320183,0.07655522,0.6268921],"study_design_scores_gemma":[0.0002935551,0.0003705498,0.0051713316,0.00022937612,0.00016298603,0.0014849433,0.00084303576,0.30888677,0.033524614,0.23565106,0.41315696,0.00022474423],"about_ca_topic_score_codex":0.004277142,"about_ca_topic_score_gemma":0.00489764,"teacher_disagreement_score":0.0055176667,"about_ca_system_score_codex":0.00078292005,"about_ca_system_score_gemma":0.0029852872,"threshold_uncertainty_score":0.029180527},"labels":[],"label_agreement":null},{"id":"W2101593757","doi":"10.1109/msr.2009.5069494","title":"On what basis to recommend: Changesets or interactions?","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Task (project management); Programmer; Software bug; Software; Recommender system; Basis (linear algebra); Information retrieval; Software engineering; Data mining; Human–computer interaction; Programming language; Engineering","score_opus":0.051171455694179364,"score_gpt":0.34193185481864025,"score_spread":0.2907603991244609,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101593757","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5891147,0.016847419,0.29046547,0.020352336,0.001878401,0.0018381964,0.011841737,0.0106807975,0.056981005],"genre_scores_gemma":[0.75796956,0.0036225156,0.22191823,0.0011575407,0.00061212794,0.00034116514,0.005638077,0.000583797,0.008156995],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9940601,0.002043482,0.0005671553,0.0012102509,0.0018348565,0.00028417338],"domain_scores_gemma":[0.96538216,0.020464774,0.0032967357,0.00422887,0.005850144,0.0007773254],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059150713,0.0012865431,0.0014132007,0.0046077785,0.0008110811,0.0034504083,0.001384163,0.002611023,0.0041560447],"category_scores_gemma":[0.055355694,0.000549492,0.0007879189,0.0044084513,0.00055745564,0.005721828,0.00061131903,0.0019119049,0.0036917776],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012427641,0.0007711433,0.12648642,0.0011572264,0.00072190375,0.00014965207,0.0010694636,0.003631219,0.0054872558,0.0022533229,0.029320523,0.82770914],"study_design_scores_gemma":[0.0013033703,0.0034055014,0.50800854,0.002219625,0.0030580491,0.0023065344,0.009060308,0.25658977,0.028470023,0.025839426,0.15849134,0.001247474],"about_ca_topic_score_codex":0.014061658,"about_ca_topic_score_gemma":0.03362673,"teacher_disagreement_score":0.014061658,"about_ca_system_score_codex":0.0007467253,"about_ca_system_score_gemma":0.0012470144,"threshold_uncertainty_score":0.031282306},"labels":[],"label_agreement":null},{"id":"W2101671828","doi":"10.1145/1275672.1275673","title":"Bridging the gap between aspect mining and refactoring","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Code refactoring; Modular programming; Computer science; Aspect-oriented programming; Bridging (networking); Object-oriented programming; Programming language; Software engineering; Identification (biology); Software","score_opus":0.048444795958268026,"score_gpt":0.3019789514318233,"score_spread":0.2535341554735553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101671828","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04933601,0.008704558,0.92567515,0.008782719,0.00014067795,0.000091379094,0.000030763113,0.0005611486,0.0066775554],"genre_scores_gemma":[0.36155516,0.0060330923,0.6282734,0.0015013502,0.00024444348,0.0001608027,0.00013916219,0.0003593364,0.0017333449],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9769116,0.012024178,0.0019563418,0.0020317843,0.0065508746,0.0005252901],"domain_scores_gemma":[0.8804908,0.08852803,0.0056941616,0.013049669,0.011269399,0.000967988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030836698,0.00072469003,0.0015430945,0.0032873,0.0008347955,0.0039003154,0.0024248313,0.0023694343,0.0010599843],"category_scores_gemma":[0.058431413,0.00094082236,0.00070140493,0.002759502,0.0025804183,0.010808508,0.0043619657,0.003157258,0.00054231286],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018751691,0.00030326305,0.011301716,0.0011504756,0.00012369362,0.0005236491,0.0054021073,0.006532802,0.008740079,0.12234488,0.002028697,0.841361],"study_design_scores_gemma":[0.00020945084,0.0007709613,0.020290736,0.0031237218,0.00022332421,0.0054617594,0.0076981448,0.2303074,0.038126886,0.56296754,0.13057895,0.00024116748],"about_ca_topic_score_codex":0.0007056997,"about_ca_topic_score_gemma":0.0008162351,"teacher_disagreement_score":0.030836698,"about_ca_system_score_codex":0.001161056,"about_ca_system_score_gemma":0.0026503874,"threshold_uncertainty_score":0.163082},"labels":[],"label_agreement":null},{"id":"W2101792768","doi":"10.5555/2337223.2337439","title":"Using the GPGPU for scaling up mining software repositories","year":2012,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; General-purpose computing on graphics processing units; Cloud computing; Eclipse; Software; Graphics; Field (mathematics); Graphics processing unit; Data science; Operating system","score_opus":0.08929231936567784,"score_gpt":0.336572874856321,"score_spread":0.24728055549064315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101792768","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49694633,0.0028671601,0.43294975,0.002744754,0.00089181965,0.0009788933,0.0015773094,0.043486886,0.017557189],"genre_scores_gemma":[0.4153733,0.0008276805,0.5780113,0.00044941113,0.0000829031,0.0005406538,0.0018861452,0.001182056,0.0016465527],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9979175,0.0005574015,0.00015011604,0.0004878811,0.000640656,0.0002464977],"domain_scores_gemma":[0.9962165,0.0011899305,0.00023222352,0.001367523,0.0007237595,0.00027022435],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016063688,0.0013709834,0.0007161392,0.002270028,0.0006909019,0.0019225696,0.0028982565,0.0008786969,0.0026475326],"category_scores_gemma":[0.012235931,0.0009013179,0.00083037966,0.004526687,0.0006825182,0.003755906,0.002649941,0.0016182285,0.0017499462],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001286261,0.00071795844,0.045874655,0.0006669984,0.00062915124,0.00064593315,0.0014290876,0.074629955,0.05069959,0.011086873,0.0518422,0.7604913],"study_design_scores_gemma":[0.0005672866,0.00089018897,0.019999713,0.00016124274,0.00029806665,0.00066377263,0.001111549,0.8574965,0.045600113,0.022393776,0.050627496,0.00019027472],"about_ca_topic_score_codex":0.009973555,"about_ca_topic_score_gemma":0.008014836,"teacher_disagreement_score":0.009973555,"about_ca_system_score_codex":0.00091852085,"about_ca_system_score_gemma":0.0014702624,"threshold_uncertainty_score":0.019831002},"labels":[],"label_agreement":null},{"id":"W2101832700","doi":"10.1016/j.scico.2009.02.007","title":"Comparison and evaluation of code clone detection techniques and tools: A qualitative approach","year":2009,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":1002,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; clone (Java method); Schema (genetic algorithms); Taxonomy (biology); Set (abstract data type); Source code; Context (archaeology); Data mining; Software engineering; Artificial intelligence; Programming language; Information retrieval","score_opus":0.09302423369983133,"score_gpt":0.41350486610888365,"score_spread":0.32048063240905234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101832700","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82606846,0.0022678804,0.12856108,0.002345588,0.00013116655,0.0075768298,0.002049899,0.00034729388,0.030651795],"genre_scores_gemma":[0.9351101,0.0007325293,0.057261102,0.0002382839,0.000020554317,0.0034368096,0.00035184,0.00011530712,0.0027335724],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.92054397,0.05109779,0.0042683785,0.0023868016,0.019684361,0.0020187385],"domain_scores_gemma":[0.6118573,0.28415164,0.01165211,0.0065709148,0.08381553,0.001952553],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.067516126,0.0009250151,0.0012527108,0.010427555,0.0033145125,0.003911766,0.00272869,0.0014926078,0.0035214317],"category_scores_gemma":[0.18859768,0.0006571646,0.0008824534,0.0061346255,0.003738722,0.0045394893,0.0033571576,0.0011769532,0.00039582246],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004674595,0.0031570676,0.060380645,0.019094188,0.00056713825,0.0008507641,0.2653497,0.0070584416,0.05505494,0.03810852,0.0057160733,0.53998786],"study_design_scores_gemma":[0.0015514122,0.017569482,0.13716246,0.010418088,0.002172018,0.001629485,0.50696975,0.038057964,0.17880596,0.03920225,0.06571417,0.00074691256],"about_ca_topic_score_codex":0.005362688,"about_ca_topic_score_gemma":0.008035923,"teacher_disagreement_score":0.067516126,"about_ca_system_score_codex":0.010356049,"about_ca_system_score_gemma":0.008320644,"threshold_uncertainty_score":0.35706365},"labels":[],"label_agreement":null},{"id":"W2102147011","doi":"10.1145/2635868.2635877","title":"Selection and presentation practices for code example summarization","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; McGill University","keywords":"Automatic summarization; Computer science; Presentation (obstetrics); Code (set theory); Selection (genetic algorithm); Source code; Code review; Think aloud protocol; Information retrieval; Software; Natural language processing; Artificial intelligence; Static program analysis; Software development; Programming language; Human–computer interaction; Usability; Set (abstract data type)","score_opus":0.04724173993157055,"score_gpt":0.32819440504035413,"score_spread":0.2809526651087836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2102147011","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3088494,0.0022971742,0.6501399,0.0051223873,0.00034068857,0.0032997166,0.0009733695,0.007093435,0.02188391],"genre_scores_gemma":[0.5769993,0.0011713906,0.4116205,0.00059225725,0.00021075421,0.002379001,0.0011612865,0.0010343122,0.004831245],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95783097,0.029807482,0.003534456,0.003175547,0.0049031815,0.0007483781],"domain_scores_gemma":[0.8375285,0.104574636,0.013121485,0.023808692,0.019383293,0.001583364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027808111,0.0008950965,0.00056830986,0.005428843,0.0021910071,0.0039781188,0.0015854624,0.0014396193,0.003973257],"category_scores_gemma":[0.16129398,0.00055407366,0.0006138921,0.0040731523,0.0016573629,0.00674538,0.00460001,0.001359085,0.0020659224],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062516716,0.00038436978,0.023487462,0.003201367,0.00012323791,0.00096763903,0.18819074,0.0013679487,0.068192914,0.013291621,0.01333725,0.68683034],"study_design_scores_gemma":[0.0003320019,0.0028451125,0.06240544,0.006076129,0.00062717585,0.00644663,0.16695377,0.03280802,0.12810896,0.061965592,0.5306218,0.0008094492],"about_ca_topic_score_codex":0.000549672,"about_ca_topic_score_gemma":0.00093857606,"teacher_disagreement_score":0.027808111,"about_ca_system_score_codex":0.0010898152,"about_ca_system_score_gemma":0.0015592381,"threshold_uncertainty_score":0.1470651},"labels":[],"label_agreement":null},{"id":"W2102394518","doi":"10.1109/ccece.2004.1349695","title":"An extension of SEMEST: the online software engineering measurement tool","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Software Engineering Process Group; Software engineering; Software measurement; Software construction; Software metric; Social software engineering; Software sizing; Software development; Personal software process; Software system; Software; Software requirements; Verification and validation; Systems engineering; Engineering","score_opus":0.030545997337515427,"score_gpt":0.2592744164570863,"score_spread":0.2287284191195709,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2102394518","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02940869,0.00028994892,0.8629609,0.0006716157,0.00022756222,0.0010446302,0.007058712,0.08947224,0.00886568],"genre_scores_gemma":[0.087681204,0.00026250695,0.89067173,0.00048990163,0.00012511811,0.0014356147,0.010443786,0.0025229196,0.0063672494],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9940578,0.0017035863,0.0010001467,0.00068822317,0.0023274687,0.00022279564],"domain_scores_gemma":[0.9673788,0.015265991,0.002986024,0.00593043,0.0077638156,0.00067490485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076924036,0.0011954488,0.0008830017,0.0036350416,0.00035816268,0.0018760123,0.0017173134,0.0010983606,0.0059594978],"category_scores_gemma":[0.035977483,0.00085597974,0.0010665043,0.0027672732,0.0004085939,0.0054072565,0.0031612301,0.0018034395,0.0033814476],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058645033,0.00092879386,0.016806452,0.0014599622,0.0002491539,0.0004473931,0.0011376478,0.011231201,0.011195451,0.014061157,0.054789577,0.8871066],"study_design_scores_gemma":[0.00054311537,0.002232856,0.04921022,0.0010708242,0.00051236927,0.0033743263,0.000569477,0.31774372,0.06365079,0.07851257,0.4820068,0.0005729417],"about_ca_topic_score_codex":0.00091930013,"about_ca_topic_score_gemma":0.0015337024,"teacher_disagreement_score":0.0076924036,"about_ca_system_score_codex":0.000498965,"about_ca_system_score_gemma":0.0018961177,"threshold_uncertainty_score":0.04068178},"labels":[],"label_agreement":null},{"id":"W2102546154","doi":"10.1109/icsm.2015.7332473","title":"Exploring API method parameter recommendations","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Java; Eclipse; Source code; Context (archaeology); Code (set theory); Property (philosophy); Programming language; Software engineering; Code generation; Operating system; Set (abstract data type)","score_opus":0.37770047939556184,"score_gpt":0.38655191054021676,"score_spread":0.00885143114465492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2102546154","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3788811,0.0058029145,0.58231,0.0018917893,0.00013266418,0.0005622241,0.0052643,0.015096309,0.010058703],"genre_scores_gemma":[0.49664888,0.0010568285,0.48808363,0.00025942366,0.00008854177,0.00030692673,0.007981462,0.0013618661,0.0042124344],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99330384,0.0019718846,0.0004118639,0.0016436069,0.0023520868,0.00031667534],"domain_scores_gemma":[0.9692415,0.020521794,0.0022619131,0.00291911,0.0046583274,0.0003973402],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004927163,0.0015059235,0.0015251704,0.0077644875,0.0010468035,0.0026593718,0.0022406816,0.0018562999,0.0022808926],"category_scores_gemma":[0.045079865,0.0008525894,0.0011928946,0.005440408,0.0004318134,0.0045647724,0.0012564571,0.0016578056,0.0012589098],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005069269,0.0007144882,0.12834916,0.0016005097,0.00045208388,0.0010899415,0.0054225395,0.013668485,0.018613469,0.0061258045,0.024165682,0.79929096],"study_design_scores_gemma":[0.00015494756,0.0005502559,0.053632338,0.0005635365,0.0006587027,0.0018486514,0.0047066268,0.7755241,0.029391242,0.016011696,0.116684124,0.0002737595],"about_ca_topic_score_codex":0.01330077,"about_ca_topic_score_gemma":0.028012034,"teacher_disagreement_score":0.01330077,"about_ca_system_score_codex":0.00096625375,"about_ca_system_score_gemma":0.0017876313,"threshold_uncertainty_score":0.0264467},"labels":[],"label_agreement":null},{"id":"W2102772625","doi":"10.1007/s10664-006-9006-4","title":"Replaying development history to assess the effectiveness of change propagation tools","year":2006,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Victoria","funders":"","keywords":"Computer science; Dependency (UML); Notice; Source code; Software engineering; Code review; Open source; Empirical research; Software development; Software; Dependency graph; Code (set theory); Change impact analysis; Data science; Static program analysis; Database; Programming language","score_opus":0.11189261393545909,"score_gpt":0.2994290649800117,"score_spread":0.18753645104455263,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2102772625","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96765745,0.00050377287,0.026642712,0.00013878722,0.00007867403,0.00014882808,0.0008356063,0.0019046917,0.002089501],"genre_scores_gemma":[0.9725568,0.00017171253,0.024805926,0.000034428063,0.000027247776,0.00009365482,0.0012570146,0.00013126002,0.00092195754],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969279,0.0013451326,0.00025872912,0.00047825338,0.0008678115,0.00012218599],"domain_scores_gemma":[0.91438156,0.06001898,0.008053178,0.009830995,0.0065079257,0.0012074624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052497666,0.0006043838,0.00052347576,0.0047651757,0.0003767566,0.00079352,0.0008947533,0.00089885364,0.0014100616],"category_scores_gemma":[0.05817827,0.00031820365,0.00039642397,0.0030500728,0.00032535515,0.0015637581,0.00058465824,0.000905736,0.00041433962],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023878426,0.0019992713,0.24384391,0.00088807626,0.0007101395,0.00036523122,0.0024334113,0.058449805,0.026634699,0.0022033236,0.0033467833,0.6567374],"study_design_scores_gemma":[0.00032037153,0.006229353,0.31685722,0.0002258722,0.0005892427,0.0009876804,0.0012124226,0.6195921,0.041462984,0.0050984863,0.007213875,0.00021039754],"about_ca_topic_score_codex":0.0024756817,"about_ca_topic_score_gemma":0.0023300548,"teacher_disagreement_score":0.0052497666,"about_ca_system_score_codex":0.00044704796,"about_ca_system_score_gemma":0.0005297733,"threshold_uncertainty_score":0.027763784},"labels":[],"label_agreement":null},{"id":"W2102866823","doi":"10.1109/cmpsac.2004.1342654","title":"An XML-based framework for language neutral program representation and generic analysis","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; XML; Programming language; Representation (politics); Document Structure Description; XML validation; Streaming XML; Efficient XML Interchange; XML Schema Editor; Program analysis; Intermediate language; World Wide Web","score_opus":0.023752088474847867,"score_gpt":0.3581027329684831,"score_spread":0.3343506444936352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2102866823","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006203936,0.00006651152,0.9945127,0.00019045454,0.000035506728,0.00009193334,0.00007377654,0.0033677341,0.0010411518],"genre_scores_gemma":[0.019677166,0.0002980886,0.97580296,0.0002570366,0.00006997778,0.00028291906,0.0006173386,0.0008372956,0.0021573284],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962132,0.001214711,0.00061964465,0.00050484197,0.0011854214,0.00026219749],"domain_scores_gemma":[0.9957098,0.0011631037,0.00042764755,0.0014197082,0.0010302323,0.00024953575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00726719,0.001139102,0.0008988767,0.0027282774,0.0013177276,0.0048122867,0.003474732,0.0023429294,0.0036621492],"category_scores_gemma":[0.00915734,0.00097496534,0.0023732586,0.0024968027,0.0023839239,0.0069153,0.0034383524,0.004386133,0.001743236],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011042877,0.0001681187,0.001019794,0.0004993483,0.00008514415,0.0005683679,0.0011916365,0.012070719,0.012644988,0.80696106,0.014818115,0.14986229],"study_design_scores_gemma":[0.000089117064,0.00020725298,0.00065497274,0.0006204804,0.00022319896,0.0015876453,0.00036226105,0.22193146,0.038687672,0.39631432,0.33911887,0.0002026753],"about_ca_topic_score_codex":0.0025904854,"about_ca_topic_score_gemma":0.0025212634,"teacher_disagreement_score":0.00726719,"about_ca_system_score_codex":0.0014421111,"about_ca_system_score_gemma":0.0026709817,"threshold_uncertainty_score":0.038433015},"labels":[],"label_agreement":null},{"id":"W2103188316","doi":"10.1145/1368088.1368154","title":"Recommending adaptive changes for framework evolution","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":156,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Code refactoring; Computer science; Software evolution; Eclipse; Software engineering; Task (project management); Simple (philosophy); Programming language; Software development; Software; Systems engineering; Engineering","score_opus":0.0661872589138669,"score_gpt":0.29611032537136367,"score_spread":0.22992306645749677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103188316","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39357546,0.002950134,0.552628,0.0018511845,0.00030428031,0.001385902,0.0028696682,0.03776035,0.0066750715],"genre_scores_gemma":[0.49028862,0.00059965963,0.50049055,0.00035310752,0.00006671179,0.00034715355,0.0043683522,0.0011258894,0.0023600091],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99635774,0.0009081543,0.00031229685,0.0011337269,0.0011045567,0.00018349254],"domain_scores_gemma":[0.9845531,0.006483599,0.0012920216,0.0029731817,0.0042473427,0.00045088198],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004157503,0.0015044486,0.001112672,0.004409793,0.0009099468,0.0016452569,0.0024374272,0.0020313128,0.0018342459],"category_scores_gemma":[0.031292222,0.0007601764,0.00088148145,0.002030673,0.00044299394,0.0025275405,0.0009634404,0.0016613338,0.00084437075],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005435124,0.00058415346,0.08458873,0.0004971447,0.00025973472,0.00043454417,0.001094157,0.027255874,0.018793788,0.0015287454,0.01837476,0.8460449],"study_design_scores_gemma":[0.00025855328,0.0006054886,0.03408696,0.0002717783,0.00058894284,0.00065319444,0.0012320145,0.88573945,0.034213692,0.004475114,0.03764295,0.0002318698],"about_ca_topic_score_codex":0.014631694,"about_ca_topic_score_gemma":0.038553756,"teacher_disagreement_score":0.014631694,"about_ca_system_score_codex":0.0011180802,"about_ca_system_score_gemma":0.0021378421,"threshold_uncertainty_score":0.029093027},"labels":[],"label_agreement":null},{"id":"W2103240721","doi":"10.1145/355045.355046","title":"Designing robust Java programs with exceptions","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Java; Computer science; Programming language; Operating system","score_opus":0.031113060094146792,"score_gpt":0.2426998679155015,"score_spread":0.2115868078213547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103240721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028437307,0.00008680968,0.96553105,0.00015559363,0.000027635815,0.00018094452,0.000021489977,0.003645749,0.0019133107],"genre_scores_gemma":[0.1971432,0.00033825674,0.7966483,0.0001784403,0.000044206765,0.0005044993,0.00015845794,0.0020196065,0.0029650277],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99770916,0.0005936192,0.0002527276,0.0004071881,0.000821639,0.0002156764],"domain_scores_gemma":[0.99548763,0.0019186537,0.0007623054,0.00094924285,0.0007164721,0.0001657399],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036220006,0.00060164125,0.00045357662,0.00052888505,0.0006866885,0.0019342335,0.0017107836,0.001093732,0.0013599194],"category_scores_gemma":[0.010800771,0.0010917687,0.00078803115,0.00034019706,0.0013144843,0.0024640474,0.0020059193,0.0015962868,0.00079447223],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027077968,0.000499648,0.009595978,0.0013584795,0.00025642235,0.0021456531,0.0049767923,0.27154434,0.15831427,0.1471478,0.007931853,0.395958],"study_design_scores_gemma":[0.00029246145,0.0006335761,0.001863161,0.00031896288,0.00037188292,0.0011865537,0.0006950938,0.70152813,0.12406142,0.059434578,0.10945337,0.00016084421],"about_ca_topic_score_codex":0.00061132107,"about_ca_topic_score_gemma":0.00086526753,"teacher_disagreement_score":0.0036220006,"about_ca_system_score_codex":0.00039312,"about_ca_system_score_gemma":0.0012591525,"threshold_uncertainty_score":0.019155145},"labels":[],"label_agreement":null},{"id":"W2103639318","doi":"10.1109/icsm.2009.5306327","title":"Playing roles in design patterns: An empirical descriptive and analytic study","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Ranking (information retrieval); Computer science; Class (philosophy); Recall; Empirical research; Zero (linguistics); Work (physics); Artificial intelligence; Psychology; Cognitive psychology; Mathematics; Statistics; Engineering","score_opus":0.07355671253990767,"score_gpt":0.3383221331416758,"score_spread":0.2647654206017681,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103639318","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98972476,0.00014843447,0.0065044207,0.0002779653,0.000003359995,0.00022741695,0.00036700364,0.000029964192,0.002716656],"genre_scores_gemma":[0.9914834,0.00012583142,0.0070836693,0.00005529983,0.000004854234,0.00029418396,0.00041765338,0.000016402832,0.00051878684],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9900831,0.004493794,0.0010242361,0.0010747662,0.002845986,0.00047809092],"domain_scores_gemma":[0.8341034,0.13268898,0.014117094,0.0056967274,0.011831138,0.0015626951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011742725,0.00042221698,0.000422635,0.0064748507,0.0019401577,0.0029321567,0.0011506879,0.0008885336,0.0015869787],"category_scores_gemma":[0.0712943,0.00058201037,0.00035658525,0.0053952886,0.003464966,0.004308583,0.00183258,0.0015249068,0.00030972395],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003684355,0.0011478339,0.7315168,0.00057038665,0.000054862834,0.0010695041,0.17460553,0.0007315456,0.0050067143,0.013697209,0.0018721406,0.069359],"study_design_scores_gemma":[0.000109334105,0.0012150947,0.5950809,0.00055473525,0.00009723765,0.0032508597,0.3463782,0.012226329,0.0066702133,0.01030534,0.02398138,0.00013033122],"about_ca_topic_score_codex":0.003333477,"about_ca_topic_score_gemma":0.0048739533,"teacher_disagreement_score":0.011742725,"about_ca_system_score_codex":0.0022531482,"about_ca_system_score_gemma":0.0016640265,"threshold_uncertainty_score":0.0621022},"labels":[],"label_agreement":null},{"id":"W2103640219","doi":"10.1109/tse.2005.28","title":"Using origin analysis to detect merging and splitting of source code entities","year":2005,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":252,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Context (archaeology); Source code; Code (set theory); Abstraction; Programming language; Software; Plan (archaeology); Software evolution; Software engineering; Database; Theoretical computer science; Software system; Data mining; Software construction","score_opus":0.023893764511130126,"score_gpt":0.2743728504647013,"score_spread":0.25047908595357116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103640219","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33771923,0.00064919505,0.637685,0.00033787618,0.000060806604,0.00022895577,0.0007312484,0.019125165,0.0034625505],"genre_scores_gemma":[0.64331436,0.0002477558,0.35109803,0.000114227776,0.000043857417,0.00012345998,0.0013101421,0.0017540533,0.0019940583],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9959111,0.0008185263,0.00041816448,0.0008737332,0.0017001289,0.00027837566],"domain_scores_gemma":[0.97344524,0.011054456,0.005633986,0.0047767805,0.004662114,0.00042734374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038887241,0.00060332543,0.0006607799,0.0065161865,0.00088273553,0.0020654867,0.0018140787,0.001224463,0.001415082],"category_scores_gemma":[0.019867152,0.0005309813,0.00090425374,0.003157291,0.0013739556,0.0036548385,0.0029741076,0.0013458348,0.0004939389],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001238193,0.00025422336,0.25722247,0.0008503843,0.0002978372,0.0040342645,0.01380366,0.010994925,0.094051644,0.02901814,0.0051191533,0.58311504],"study_design_scores_gemma":[0.0002247944,0.0006805702,0.107957356,0.00026817364,0.0007655524,0.0053645596,0.0032623126,0.44403163,0.34021935,0.038992427,0.05782363,0.00040969584],"about_ca_topic_score_codex":0.003101582,"about_ca_topic_score_gemma":0.0029623036,"teacher_disagreement_score":0.0065161865,"about_ca_system_score_codex":0.0009171635,"about_ca_system_score_gemma":0.0011919173,"threshold_uncertainty_score":0.020565808},"labels":[],"label_agreement":null},{"id":"W2103709708","doi":"10.1109/compsac.2011.67","title":"A Semi-automated Decision Support Tool for Requirements Trade-Off Analysis","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Requirements analysis; Usability; Computer science; Requirements management; Process (computing); Measure (data warehouse); Risk analysis (engineering); Requirements engineering; Requirements elicitation; Business requirements; Functional requirement; Requirement; Decision support system; User requirements document; Non-functional requirement; Requirement prioritization; Systems engineering; Software engineering; Database; Data mining; Engineering; Software; Work in process; Human–computer interaction; Operations management; Business process; Software development","score_opus":0.04750607050300762,"score_gpt":0.31792560163848926,"score_spread":0.27041953113548167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103709708","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020535353,0.00003156762,0.9825867,0.0000848866,0.000017009943,0.00042792512,0.00058406795,0.012805901,0.001408494],"genre_scores_gemma":[0.013059035,0.00002887722,0.98456407,0.000035089306,0.000008297483,0.0007304328,0.0006158491,0.00035873376,0.00059964263],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99397177,0.0028416857,0.00074158935,0.00055869675,0.0017239543,0.00016228949],"domain_scores_gemma":[0.96690416,0.026420001,0.0013825853,0.0021778517,0.0026853539,0.00043007662],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009558713,0.0026101724,0.0018208545,0.00422296,0.0014568599,0.0031160037,0.0027270138,0.0017257164,0.02539986],"category_scores_gemma":[0.026247574,0.0013432116,0.0017074724,0.0025441404,0.0009682715,0.0026044422,0.0031881428,0.0022508865,0.0064587756],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008339976,0.00077462057,0.0017051138,0.001665912,0.00024019546,0.0009053607,0.001383781,0.09782428,0.023349212,0.041520234,0.040863466,0.78893393],"study_design_scores_gemma":[0.00044812215,0.00023471382,0.0006637302,0.00025559688,0.00007037064,0.00044736444,0.00023749482,0.9009338,0.012468857,0.042397954,0.041708462,0.00013346235],"about_ca_topic_score_codex":0.0019715775,"about_ca_topic_score_gemma":0.0028283899,"teacher_disagreement_score":0.02539986,"about_ca_system_score_codex":0.0013369597,"about_ca_system_score_gemma":0.0034106434,"threshold_uncertainty_score":0.08497101},"labels":[],"label_agreement":null},{"id":"W2103715428","doi":"10.1109/nafips.2007.383813","title":"Applying Novel Resampling Strategies To Software Defect Prediction","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":160,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Machine learning; Computer science; Resampling; Artificial intelligence; Software bug; Benchmark (surveying); Software; Software quality; Data mining; Class (philosophy); Skewness; Software metric; Classifier (UML); Software development","score_opus":0.03284797929501758,"score_gpt":0.29673545371281235,"score_spread":0.26388747441779475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103715428","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14893064,0.0005510284,0.84835154,0.00030477787,0.000089731344,0.00015093615,0.000107975015,0.0008880566,0.00062523416],"genre_scores_gemma":[0.70094395,0.00021524943,0.29695836,0.00020735901,0.0001339974,0.00018647358,0.0006263904,0.00006966494,0.00065852946],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978688,0.0011364521,0.00012799197,0.00029993537,0.00045864473,0.000108168235],"domain_scores_gemma":[0.99057627,0.005593548,0.00074192596,0.001360603,0.0015364672,0.00019127881],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057376386,0.000747412,0.0011126306,0.0018431935,0.0004920651,0.0005531671,0.0010990389,0.0008829386,0.00038429335],"category_scores_gemma":[0.017439077,0.00032068355,0.00077472,0.0007380631,0.00062027865,0.0011038696,0.0008338717,0.00090264494,0.00020253207],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007641438,0.0006585327,0.022252586,0.00014432904,0.0003605242,0.00029814104,0.00044848627,0.4477357,0.023759829,0.006994079,0.003748628,0.4928349],"study_design_scores_gemma":[0.000023698172,0.00013993599,0.0020023205,0.0000073429615,0.000018426537,0.000059673857,0.000028739863,0.9891054,0.0048020137,0.0032542255,0.0005446781,0.000013517472],"about_ca_topic_score_codex":0.0026326554,"about_ca_topic_score_gemma":0.0042280387,"teacher_disagreement_score":0.0057376386,"about_ca_system_score_codex":0.0004960052,"about_ca_system_score_gemma":0.00060402683,"threshold_uncertainty_score":0.03034389},"labels":[],"label_agreement":null},{"id":"W2103736207","doi":"10.1109/wcre.2003.1287246","title":"Predicting maintainability with object-oriented metrics -an empirical comparison","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":129,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Maintainability; Cohesion (chemistry); Computer science; Object-oriented programming; Software metric; Software quality; Inheritance (genetic algorithm); Software evolution; Software measurement; Software sizing; Software; Empirical research; Software maintenance; Software engineering; Software system; Reliability engineering; Software development; Software construction; Programming language; Engineering; Statistics; Mathematics","score_opus":0.02824108292722796,"score_gpt":0.3273465326844309,"score_spread":0.299105449757203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103736207","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99211884,0.00091813353,0.0048567806,0.0001131833,0.000011365387,0.000032806864,0.00062260084,0.000063860076,0.0012624699],"genre_scores_gemma":[0.99633384,0.00028382064,0.0021007457,0.000016796383,0.000019455105,0.00002668035,0.001014825,0.000023154484,0.00018062067],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9921182,0.00375746,0.00074599276,0.00076846575,0.0024054085,0.0002045225],"domain_scores_gemma":[0.8321984,0.135961,0.013427099,0.00898994,0.008197097,0.0012265459],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014748659,0.0006647564,0.00039934463,0.005341599,0.00025569883,0.0013476214,0.00092191267,0.0009013544,0.0010599662],"category_scores_gemma":[0.095066644,0.0002518986,0.0005119994,0.00430269,0.0006983192,0.003180096,0.0009055423,0.0006662157,0.00045130332],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001074624,0.00017613592,0.97182775,0.00006293624,0.00023499361,0.00004391563,0.00020284472,0.0034783387,0.00023741291,0.00019574344,0.00031364633,0.023118945],"study_design_scores_gemma":[0.000026962778,0.0006731396,0.9551182,0.000053459415,0.00012323656,0.0003181532,0.00048272958,0.03983876,0.0011583648,0.0009998145,0.0011782177,0.000028908691],"about_ca_topic_score_codex":0.001967829,"about_ca_topic_score_gemma":0.0026689603,"teacher_disagreement_score":0.014748659,"about_ca_system_score_codex":0.00040598403,"about_ca_system_score_gemma":0.00023561002,"threshold_uncertainty_score":0.077999294},"labels":[],"label_agreement":null},{"id":"W2103799318","doi":"10.1109/icsm.2008.4658107","title":"A domain-customizable SVG-based graph editor for software visualizations","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Scalable Vector Graphics; Computer science; JavaScript; sync; Scalability; Graphics; Graph; Software; Programming language; World Wide Web; Domain (mathematical analysis); Vector graphics; Computer graphics (images); Database; Theoretical computer science","score_opus":0.01905190899150036,"score_gpt":0.27510934385481095,"score_spread":0.25605743486331056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103799318","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027389438,0.000074773256,0.8784202,0.00010073219,0.000115392926,0.00007972039,0.0012111893,0.115525976,0.0017331982],"genre_scores_gemma":[0.040212203,0.00026255532,0.9240816,0.00016714803,0.00007246989,0.000252418,0.0037805692,0.027230186,0.003940791],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.999246,0.00015294095,0.00011761603,0.00014770882,0.00028627738,0.000049504764],"domain_scores_gemma":[0.99624324,0.0018477497,0.00016516974,0.000878989,0.0006453854,0.00021950994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014018334,0.0010715945,0.0008680757,0.0011676953,0.00027848195,0.0014769265,0.0023499685,0.0006781426,0.011010284],"category_scores_gemma":[0.0057275794,0.00070127565,0.0009687052,0.00086679374,0.00036907135,0.0018428874,0.0014122584,0.0019316798,0.003945761],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006052047,0.00023667066,0.001893154,0.0015791178,0.00021643045,0.0015094784,0.00074918766,0.035395753,0.15063962,0.034887545,0.1663904,0.6058975],"study_design_scores_gemma":[0.00043715784,0.00017476502,0.0014994334,0.00019979493,0.00013029315,0.0016110622,0.00014223873,0.39682028,0.14792217,0.03828303,0.41252953,0.00025027106],"about_ca_topic_score_codex":0.00098641,"about_ca_topic_score_gemma":0.0016384817,"teacher_disagreement_score":0.011010284,"about_ca_system_score_codex":0.00029153845,"about_ca_system_score_gemma":0.0007161641,"threshold_uncertainty_score":0.036833107},"labels":[],"label_agreement":null},{"id":"W2103897902","doi":"10.1109/wpc.2001.921708","title":"Software visualization tools: survey and analysis","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":112,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Visualization; Program comprehension; Data science; Software; Software engineering; Data visualization; Software inspection; Human–computer interaction; Software development; Software quality; Software system; Data mining; Programming language","score_opus":0.055589986230228705,"score_gpt":0.2963531649433472,"score_spread":0.24076317871311848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103897902","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86472505,0.06738926,0.025261732,0.0029239194,0.0002287602,0.0024552487,0.012777112,0.0008322337,0.02340668],"genre_scores_gemma":[0.8871347,0.05921896,0.02901825,0.0011962544,0.00020708035,0.0032944689,0.012295747,0.00028555293,0.0073490366],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98332006,0.00593199,0.0028634996,0.0010534558,0.0058788043,0.00095212285],"domain_scores_gemma":[0.9023125,0.06504166,0.006407146,0.002139312,0.022433894,0.0016654703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013729899,0.00049305503,0.0011422554,0.016414791,0.0007804757,0.0021160005,0.00072081364,0.0010419049,0.0027587267],"category_scores_gemma":[0.048632562,0.00052744243,0.0007566401,0.017858244,0.00060626346,0.0037038112,0.0015280531,0.0009165477,0.0012423553],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007453647,0.0009345678,0.20788151,0.010564818,0.00026660238,0.0003995826,0.018700892,0.0009970345,0.0035735818,0.0021038328,0.024629777,0.72920233],"study_design_scores_gemma":[0.00010372378,0.0036162043,0.6433752,0.008234491,0.0004175451,0.003287559,0.04733862,0.0039818143,0.00565102,0.0017267739,0.2820582,0.0002088063],"about_ca_topic_score_codex":0.0014615762,"about_ca_topic_score_gemma":0.001886785,"teacher_disagreement_score":0.016414791,"about_ca_system_score_codex":0.0010530755,"about_ca_system_score_gemma":0.0019497706,"threshold_uncertainty_score":0.07261151},"labels":[],"label_agreement":null},{"id":"W2103920581","doi":"10.1109/achi.2008.33","title":"Examining Programmer's Cognitive Skills Using Regular Language","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Saint Mary's University; Dalhousie University","funders":"","keywords":"Computer science; Regular expression; Programmer; Notation; Pattern matching; Cognition; Task (project management); Programming language; Alternation (linguistics); Matching (statistics); Natural language; Completeness (order theory); Artificial intelligence; Natural language processing; Mathematics; Arithmetic; Linguistics; Psychology","score_opus":0.04375681411313004,"score_gpt":0.29368222455851695,"score_spread":0.2499254104453869,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103920581","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99291754,0.00004334666,0.0016023361,0.00007103624,0.0000040508685,0.000020503983,0.00004831599,0.000067001085,0.005225959],"genre_scores_gemma":[0.99585646,0.00007472316,0.0021296712,0.000045369678,0.0000050700223,0.000023576795,0.0000941906,0.000017075121,0.0017537304],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9987143,0.00036552668,0.00015132992,0.00024389414,0.0003993984,0.00012554045],"domain_scores_gemma":[0.9459903,0.03601732,0.008356474,0.0031134663,0.005195605,0.0013269054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002046203,0.0003364623,0.00017159298,0.0012117513,0.00019024564,0.0014924756,0.00043116004,0.0004888475,0.0029738673],"category_scores_gemma":[0.03991193,0.0002099151,0.0002146105,0.0004695025,0.0005058759,0.0014687286,0.0006491158,0.00071442215,0.000680856],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004905182,0.0013370635,0.7463362,0.00045408594,0.00015823019,0.000381191,0.026562687,0.0042232745,0.023060832,0.0023410122,0.0023937398,0.19226117],"study_design_scores_gemma":[0.00007519095,0.0018148517,0.9566978,0.0001159331,0.000094378236,0.0010874281,0.008897758,0.012383681,0.008370896,0.004610467,0.0057572555,0.000094403265],"about_ca_topic_score_codex":0.0012592191,"about_ca_topic_score_gemma":0.0018644708,"teacher_disagreement_score":0.0029738673,"about_ca_system_score_codex":0.00019494562,"about_ca_system_score_gemma":0.00042895428,"threshold_uncertainty_score":0.010821462},"labels":[],"label_agreement":null},{"id":"W2104040631","doi":"10.1109/wcre.2006.48","title":"Refactoring Detection based on UMLDiff Change-Facts Queries","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":102,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Code refactoring; Computer science; Software evolution; Software engineering; Software; Software system; Object-oriented programming; Software development; Software quality; Software maintenance; Software metric; Programming language; Software construction","score_opus":0.02923127407523809,"score_gpt":0.2539968840884176,"score_spread":0.2247656100131795,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104040631","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15759693,0.0005800717,0.81998074,0.0007001132,0.000059421887,0.00076034846,0.0017297218,0.016371738,0.0022210067],"genre_scores_gemma":[0.3823948,0.00019507729,0.6127654,0.00018615612,0.000047299556,0.00029288285,0.0026162944,0.0005581769,0.0009439193],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98913246,0.0029625616,0.0015133109,0.0017182068,0.0042732367,0.00040030797],"domain_scores_gemma":[0.9416449,0.03784346,0.0060840654,0.0057342346,0.007987665,0.00070560945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007769408,0.0014240775,0.0016208661,0.008310433,0.00078965863,0.002335068,0.0026152446,0.0028273962,0.0017283075],"category_scores_gemma":[0.05265898,0.0007416161,0.0008312065,0.0032843784,0.0009340673,0.0037617574,0.0017979104,0.0012551394,0.0008464108],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010593728,0.00081076886,0.1080538,0.0009892443,0.00023155219,0.0022004496,0.0041077794,0.032758325,0.07688488,0.014292834,0.008387881,0.75022316],"study_design_scores_gemma":[0.00019084848,0.000336835,0.03011444,0.000089927016,0.00014609881,0.0016391583,0.00068556686,0.88415563,0.060312655,0.010413716,0.01175265,0.00016252433],"about_ca_topic_score_codex":0.0056483434,"about_ca_topic_score_gemma":0.0071910066,"teacher_disagreement_score":0.008310433,"about_ca_system_score_codex":0.0011326758,"about_ca_system_score_gemma":0.001467692,"threshold_uncertainty_score":0.041089058},"labels":[],"label_agreement":null},{"id":"W2104086546","doi":"10.1145/2020390.2020401","title":"Studying the fix-time for bugs in large open source projects","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":118,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software bug; Open source; Computer science; Software; Open source software; Computer security; Operating system","score_opus":0.0846688826073423,"score_gpt":0.3074386331270086,"score_spread":0.2227697505196663,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104086546","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99290586,0.0010066631,0.0037402692,0.00022774203,0.0000144741,0.000046919704,0.00038840107,0.00006894547,0.001600799],"genre_scores_gemma":[0.9942167,0.00040845916,0.0036988817,0.000029907102,0.000026097605,0.00008399613,0.00087940897,0.000043578322,0.00061298127],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9937482,0.0018987893,0.0006364449,0.0011150357,0.0021866565,0.00041496695],"domain_scores_gemma":[0.7429068,0.1668162,0.057649452,0.008891018,0.016590865,0.0071456903],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0097437715,0.00049997185,0.00038848398,0.0053656446,0.0009924414,0.0015372437,0.0010807685,0.001126037,0.0031379964],"category_scores_gemma":[0.13324301,0.000559017,0.00062090263,0.0040009143,0.0008078773,0.0033100448,0.0018936285,0.0016379271,0.00071214937],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041770437,0.0004363791,0.9125527,0.00034697173,0.00026851663,0.00022150899,0.0030031996,0.0037821045,0.0021876474,0.001596838,0.0013465133,0.07383986],"study_design_scores_gemma":[0.000038375736,0.000557012,0.98692465,0.000119419565,0.00008622338,0.00036884213,0.0014020344,0.0052743205,0.0007605117,0.0024777479,0.0019420523,0.000048820162],"about_ca_topic_score_codex":0.005182982,"about_ca_topic_score_gemma":0.006817518,"teacher_disagreement_score":0.0097437715,"about_ca_system_score_codex":0.0014911927,"about_ca_system_score_gemma":0.0011939376,"threshold_uncertainty_score":0.0515306},"labels":[],"label_agreement":null},{"id":"W2104193835","doi":"10.1109/vlhcc.2006.49","title":"Using Visual Momentum to Explain Disorientation in the Eclipse IDE","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Eclipse; Computer science; Context (archaeology); Java; Code (set theory); Field (mathematics); Software; Momentum (technical analysis); Human–computer interaction; Programming language; Software engineering; History","score_opus":0.02802982897556398,"score_gpt":0.3297490925436858,"score_spread":0.3017192635681218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104193835","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98033303,0.00017579996,0.015130117,0.00030756788,0.000013409167,0.000032773776,0.000033134307,0.00008107396,0.0038930648],"genre_scores_gemma":[0.9982332,0.00004800755,0.001490194,0.000019540532,0.0000029864464,0.000010915844,0.000022381502,0.000013559644,0.00015918665],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.998035,0.0007098578,0.00017271223,0.00019419707,0.00055976375,0.00032838035],"domain_scores_gemma":[0.9693132,0.019709991,0.007097627,0.0017062009,0.0016034099,0.00056963554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027251355,0.00061245216,0.00025947645,0.001602034,0.00077943964,0.0014922027,0.0005104452,0.00061385677,0.0012072875],"category_scores_gemma":[0.028739111,0.0003985091,0.00029306393,0.000951023,0.002298241,0.0033721735,0.0018559095,0.0011997637,0.000110737616],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024974896,0.0006735364,0.676662,0.0004963525,0.00010091515,0.0029032482,0.10288856,0.010517215,0.03894394,0.020337945,0.0012168855,0.14276189],"study_design_scores_gemma":[0.00023906371,0.0013690398,0.80825305,0.00028235398,0.0001430416,0.0025842285,0.06607652,0.048847593,0.015660005,0.044453107,0.011812756,0.0002791627],"about_ca_topic_score_codex":0.0061223693,"about_ca_topic_score_gemma":0.006806968,"teacher_disagreement_score":0.0061223693,"about_ca_system_score_codex":0.0012768132,"about_ca_system_score_gemma":0.0008244696,"threshold_uncertainty_score":0.014412105},"labels":[],"label_agreement":null},{"id":"W2104572159","doi":"10.1109/sew.2005.42","title":"Supporting Software Release Planning Decisions for Evolving Systems","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Software release life cycle; Resource (disambiguation); Risk analysis (engineering); Key (lock); Software; Software system; Enterprise resource planning; Feature (linguistics); Resource planning; Systems engineering; Process management; Software engineering; Knowledge management; Engineering; Software construction; Computer security; Business","score_opus":0.028289096133592335,"score_gpt":0.313308007261154,"score_spread":0.28501891112756167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104572159","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24469654,0.00040507125,0.7446839,0.0006252733,0.0000382291,0.0004982689,0.00022311672,0.0035114533,0.0053181564],"genre_scores_gemma":[0.69099146,0.00024745116,0.3072306,0.000042579148,0.000021323436,0.00014308337,0.0002963792,0.00016118791,0.000865936],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99672407,0.0015637651,0.00022196067,0.0003480503,0.0008973989,0.00024475355],"domain_scores_gemma":[0.976409,0.017034711,0.0029405889,0.0013339074,0.0017638587,0.0005180169],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064825243,0.0010416087,0.00077519484,0.0014177731,0.0010377201,0.0024245903,0.001120845,0.0008446303,0.0018895009],"category_scores_gemma":[0.02409202,0.00081090676,0.00060137466,0.00074023433,0.0006602517,0.0019510151,0.0012883001,0.0011462043,0.00032327525],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005835312,0.0003852891,0.010047714,0.00031856896,0.00010939807,0.0005362339,0.001027519,0.61047477,0.009777106,0.012109091,0.0028444778,0.3517863],"study_design_scores_gemma":[0.00007722201,0.00016178766,0.0013465548,0.000029069734,0.000035392153,0.00009393934,0.00025952948,0.9868025,0.0052610342,0.0040739286,0.0018314593,0.000027557313],"about_ca_topic_score_codex":0.0058163987,"about_ca_topic_score_gemma":0.008336646,"teacher_disagreement_score":0.0064825243,"about_ca_system_score_codex":0.0009829492,"about_ca_system_score_gemma":0.0022172716,"threshold_uncertainty_score":0.03428328},"labels":[],"label_agreement":null},{"id":"W2104607014","doi":"10.1109/icsm.2012.6405278","title":"Relating requirements to implementation via topic analysis: Do topics extracted from requirements make sense to managers and developers?","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Traceability; Documentation; Relevance (law); Perception; Topic model; Control (management); Requirements analysis; Requirements traceability; Requirements engineering; Requirements elicitation; Software engineering; Software; Information retrieval; Requirement; Artificial intelligence; Programming language","score_opus":0.04333193516344207,"score_gpt":0.3415317457067007,"score_spread":0.2981998105432586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104607014","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75626194,0.0011210329,0.23436491,0.0011699552,0.000071904295,0.00045188676,0.0009182695,0.00051229977,0.0051277727],"genre_scores_gemma":[0.952178,0.0003187427,0.04534004,0.00009499339,0.000052231742,0.00031381036,0.0011276598,0.000098130105,0.00047638116],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.98782396,0.0076998,0.00086831267,0.0016716754,0.0014818268,0.00045431659],"domain_scores_gemma":[0.8115293,0.1582367,0.012683628,0.006786059,0.009811144,0.0009530981],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018605826,0.00059647515,0.00066565187,0.00613895,0.0010406743,0.0038254645,0.0008267644,0.0011114284,0.00097563857],"category_scores_gemma":[0.10910126,0.0004638532,0.0010060889,0.005195888,0.0011014637,0.005910632,0.0018347716,0.0016534486,0.0004413398],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013254143,0.0004157903,0.45395297,0.0017118276,0.00060663745,0.00032260318,0.068699256,0.011457945,0.013939249,0.009453299,0.0035774892,0.43453756],"study_design_scores_gemma":[0.00017085507,0.00078903005,0.63184994,0.0007753848,0.00077094627,0.00082414574,0.051594656,0.23726125,0.013431372,0.042284258,0.019807898,0.00044025312],"about_ca_topic_score_codex":0.005637368,"about_ca_topic_score_gemma":0.004660723,"teacher_disagreement_score":0.018605826,"about_ca_system_score_codex":0.0014977868,"about_ca_system_score_gemma":0.0012967649,"threshold_uncertainty_score":0.09839821},"labels":[],"label_agreement":null},{"id":"W2104609444","doi":"10.5555/2664398.2664404","title":"Java bytecode clone detection via relaxation on code fingerprint and semantic web reasoning","year":2012,"lang":"en","type":"article","venue":"International Workshop on Software Clones","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Concordia University","funders":"","keywords":"Bytecode; Computer science; Java bytecode; Java; Source code; Programming language; Theoretical computer science; Artificial intelligence; Java applet; Java annotation","score_opus":0.018768790849269648,"score_gpt":0.2805580122239025,"score_spread":0.26178922137463284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104609444","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1335297,0.000320334,0.858931,0.00016908489,0.00001539871,0.00019032513,0.00013165432,0.004646373,0.002066062],"genre_scores_gemma":[0.5517334,0.00015972728,0.44618252,0.000092117094,0.000013449174,0.00013103434,0.0004250698,0.00021935436,0.001043379],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958646,0.00072000106,0.00030369931,0.0008185444,0.002044971,0.00024814677],"domain_scores_gemma":[0.9924436,0.003130745,0.0013012738,0.001322559,0.0016548139,0.0001469828],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021146475,0.00066170696,0.0010676999,0.0060265,0.0007218988,0.0019559714,0.0015471072,0.0011513485,0.0011191762],"category_scores_gemma":[0.014268836,0.0004627767,0.0013100567,0.003191207,0.001366529,0.0037486227,0.0022394964,0.0008661783,0.00035112698],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045376105,0.000382221,0.02948972,0.0004915028,0.00016834792,0.0006506966,0.0014804525,0.043660976,0.07660129,0.023656318,0.0016803825,0.82128435],"study_design_scores_gemma":[0.00004436434,0.00019916479,0.012049282,0.000063453415,0.00014069595,0.0009393359,0.00054847024,0.8845475,0.06420424,0.032899093,0.0042937496,0.00007054943],"about_ca_topic_score_codex":0.0045274966,"about_ca_topic_score_gemma":0.003804771,"teacher_disagreement_score":0.0060265,"about_ca_system_score_codex":0.0011399785,"about_ca_system_score_gemma":0.001523024,"threshold_uncertainty_score":0.0111835},"labels":[],"label_agreement":null},{"id":"W2104613753","doi":"10.1109/icsm.2007.4362641","title":"Release Pattern Discovery: A Case Study of Database Systems","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Documentation; Database; Behavioral pattern; Data mining; Software engineering; Programming language","score_opus":0.029752065513547945,"score_gpt":0.30043389423696737,"score_spread":0.27068182872341945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104613753","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99090064,0.00023936186,0.0071286783,0.00045633808,0.00000921616,0.00015028892,0.0003032025,0.00008626202,0.0007260224],"genre_scores_gemma":[0.95921814,0.0003955863,0.037743855,0.00016897643,0.000026815187,0.00015341683,0.00090856914,0.00008520289,0.0012994962],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9909616,0.003933188,0.0008860238,0.0011832424,0.0024661832,0.00056971645],"domain_scores_gemma":[0.91956383,0.06107631,0.006578666,0.0062820963,0.004882718,0.0016163543],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066799745,0.00055878185,0.00055971387,0.003773382,0.0024091217,0.0015354026,0.002187603,0.0022981097,0.00054568815],"category_scores_gemma":[0.036051393,0.000729697,0.0009566507,0.004089439,0.001392947,0.002195274,0.0016081417,0.0016106325,0.00019076712],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081693503,0.0042869984,0.66336507,0.0015854092,0.000514496,0.039389394,0.07517806,0.013228446,0.021549122,0.0050698486,0.0059215557,0.16909467],"study_design_scores_gemma":[0.0005705441,0.0028609158,0.69763315,0.00046803846,0.00049417745,0.046529908,0.061807144,0.10958776,0.037805583,0.007122242,0.034755323,0.00036530112],"about_ca_topic_score_codex":0.011110135,"about_ca_topic_score_gemma":0.021372106,"teacher_disagreement_score":0.011110135,"about_ca_system_score_codex":0.0010178307,"about_ca_system_score_gemma":0.0010412169,"threshold_uncertainty_score":0.035327494},"labels":[],"label_agreement":null},{"id":"W2104844656","doi":"10.1145/1806799.1806854","title":"Awareness 2.0","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":117,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Task (project management); Prioritization; Process (computing); Process management; Knowledge management; Event (particle physics); Software project management; Project management; Project team; Software development process; Software development; Software; Systems engineering; Engineering; Software construction","score_opus":0.01608775821799958,"score_gpt":0.2825527402977476,"score_spread":0.266464982079748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104844656","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046979673,0.0019051878,0.38742113,0.00079340563,0.00053157937,0.00093296403,0.008303749,0.52688986,0.06852416],"genre_scores_gemma":[0.11394516,0.00379263,0.57658017,0.0024980481,0.00059724745,0.003057223,0.051956102,0.13887818,0.10869531],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973507,0.00049828424,0.00032117774,0.0005818121,0.00091242715,0.00033572322],"domain_scores_gemma":[0.9952342,0.0014916385,0.00042122317,0.0016880016,0.00074878003,0.00041623798],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003956191,0.0018106246,0.0011673316,0.002516847,0.00066052267,0.005494796,0.0032368738,0.0018997235,0.031578336],"category_scores_gemma":[0.011792615,0.0022290545,0.0018995671,0.001451113,0.000602444,0.005608372,0.0053941067,0.0030850617,0.032261696],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013291924,0.00032866202,0.0045574103,0.0028085262,0.00039699578,0.0006393676,0.002050464,0.0043024374,0.016379762,0.07276577,0.30891246,0.58552897],"study_design_scores_gemma":[0.00019124393,0.00021155273,0.0031339664,0.00056879054,0.00019832057,0.0011483349,0.00015736668,0.014176171,0.01906289,0.047737624,0.9131726,0.00024111029],"about_ca_topic_score_codex":0.001366285,"about_ca_topic_score_gemma":0.0009003326,"teacher_disagreement_score":0.031578336,"about_ca_system_score_codex":0.0005841308,"about_ca_system_score_gemma":0.0019437799,"threshold_uncertainty_score":0.105640054},"labels":[],"label_agreement":null},{"id":"W2104872620","doi":"10.1145/337180.343189","title":"Characterizing implicit information during peer review meetings","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Division of Materials Research; Polytechnique Montréal; Natural Sciences and Engineering Research Council of Canada; Institut national de recherche en informatique et en automatique (INRIA)","keywords":"Computer science; World Wide Web; Information retrieval; Data science; Knowledge management","score_opus":0.012121023261218807,"score_gpt":0.25834028478096016,"score_spread":0.24621926151974136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104872620","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92702544,0.0018854797,0.05128766,0.0023757163,0.00021074894,0.00066479575,0.0006043425,0.0004219441,0.0155239515],"genre_scores_gemma":[0.9871651,0.00025215495,0.009550419,0.00012866416,0.00027458539,0.00027787723,0.00046228102,0.000065892564,0.0018231255],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.791235,0.10759141,0.017332533,0.011757086,0.06714702,0.0049369554],"domain_scores_gemma":[0.2568235,0.51173013,0.11581112,0.046043124,0.06231842,0.0072737546],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06496204,0.0007026818,0.0014167286,0.008415703,0.0030629393,0.0060714735,0.0027114013,0.0027722758,0.0019968078],"category_scores_gemma":[0.44384828,0.0007079341,0.0005689814,0.005339727,0.0018101437,0.008433352,0.005117459,0.002401691,0.0012345909],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024633147,0.0010083936,0.43455467,0.0018524904,0.0007661484,0.0021490208,0.068447195,0.0049444567,0.024085255,0.011905943,0.005523636,0.4422995],"study_design_scores_gemma":[0.0002746053,0.0035014576,0.7752032,0.0009997295,0.00071978156,0.003796136,0.039494045,0.05983965,0.026910886,0.04018833,0.048281483,0.0007906875],"about_ca_topic_score_codex":0.0010989079,"about_ca_topic_score_gemma":0.0013138464,"teacher_disagreement_score":0.935038,"about_ca_system_score_codex":0.0020770144,"about_ca_system_score_gemma":0.0033453656,"threshold_uncertainty_score":0.34355617},"labels":[],"label_agreement":null},{"id":"W2104890846","doi":"10.1145/1370175.1370211","title":"CCVisu","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Software; VRML; File format; Software system; Graph; Component-based software engineering; Theoretical computer science; Graph rewriting; Programming language; Operating system; The Internet","score_opus":0.029445631274286448,"score_gpt":0.26119623602949943,"score_spread":0.23175060475521297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104890846","genre_codex":"other","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041308613,0.001603653,0.18271357,0.0018560047,0.0022543909,0.00052741915,0.029426716,0.35868,0.4188074],"genre_scores_gemma":[0.050263744,0.0022195906,0.20011975,0.0020771257,0.0011291996,0.001724231,0.14672542,0.12685938,0.46888155],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99790514,0.0003415682,0.00010268166,0.00044336636,0.00097633526,0.00023095003],"domain_scores_gemma":[0.9956234,0.00082394725,0.00017305108,0.0015000501,0.0013253218,0.00055432617],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020307398,0.0017155103,0.0012856018,0.004055478,0.0016414899,0.0063786274,0.004401657,0.0027022234,0.3224954],"category_scores_gemma":[0.009858152,0.00090351,0.0012182315,0.003783029,0.0009571243,0.005479893,0.0046045976,0.0021476848,0.21943453],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030850168,0.00008987994,0.0006967592,0.00042327738,0.000033590546,0.00017913617,0.0002444763,0.0009911505,0.00253678,0.024983494,0.70443785,0.26507518],"study_design_scores_gemma":[0.000057762965,0.00003519139,0.00035954252,0.000109727705,0.000009258403,0.00017867064,0.00005053531,0.004295696,0.0025297217,0.011136472,0.981203,0.000034368848],"about_ca_topic_score_codex":0.0044072126,"about_ca_topic_score_gemma":0.003540666,"teacher_disagreement_score":0.3224954,"about_ca_system_score_codex":0.0018349261,"about_ca_system_score_gemma":0.0028084982,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2104905305","doi":"10.1109/wpc.2000.852495","title":"On the stability of software clustering algorithms","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Toronto","funders":"","keywords":"Decomposition; Computer science; Cluster analysis; Software; Stability (learning theory); Software system; Algorithm; Property (philosophy); Data mining; Theoretical computer science; Artificial intelligence; Programming language; Machine learning","score_opus":0.04764062874470087,"score_gpt":0.2557499053697583,"score_spread":0.2081092766250574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104905305","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0823679,0.0025979516,0.9047686,0.0012258567,0.00014835873,0.00015627526,0.00016600722,0.0008107605,0.007758296],"genre_scores_gemma":[0.7299934,0.0020488813,0.26129442,0.00048141752,0.00048446326,0.00042427034,0.00070646673,0.0006992757,0.0038674339],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9897552,0.004779485,0.0006307645,0.0016995694,0.002541626,0.0005933359],"domain_scores_gemma":[0.8818249,0.09426151,0.0047030337,0.006507348,0.011519698,0.0011835659],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012165602,0.0013715429,0.0020068742,0.0041522216,0.0030329982,0.0037126197,0.002038459,0.0023381985,0.0029423258],"category_scores_gemma":[0.105371624,0.0008353947,0.0014337628,0.0036585163,0.004743189,0.0054937825,0.0038907714,0.0030051344,0.0011817102],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012920477,0.00022396316,0.010662732,0.00045175597,0.00030491193,0.00021185871,0.0013538373,0.54129726,0.00801152,0.26534337,0.005876936,0.16496973],"study_design_scores_gemma":[0.00006435025,0.00012496435,0.0009019547,0.000060159244,0.000036142,0.00009953785,0.000107454296,0.81585175,0.0025309036,0.17857903,0.0016094513,0.000034234352],"about_ca_topic_score_codex":0.0035165008,"about_ca_topic_score_gemma":0.0015090959,"teacher_disagreement_score":0.012165602,"about_ca_system_score_codex":0.0024256501,"about_ca_system_score_gemma":0.0018863919,"threshold_uncertainty_score":0.064338624},"labels":[],"label_agreement":null},{"id":"W2104944621","doi":"10.1109/icsm.2009.5306366","title":"Understanding source package organization using the hybrid model","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Namespace; Cohesion (chemistry); Reuse; Object-oriented programming; Programming language; Software engineering; Software design pattern; Inheritance (genetic algorithm); Design pattern; Construct (python library); Set (abstract data type); Software; Database","score_opus":0.08687297082512917,"score_gpt":0.27352525049090587,"score_spread":0.1866522796657767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104944621","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06291983,0.00011830285,0.92672807,0.0007111384,0.000010989386,0.0000888737,0.0001108816,0.000700182,0.008611713],"genre_scores_gemma":[0.53367317,0.00028622386,0.459879,0.00017048165,0.000024250867,0.00036627866,0.00037570743,0.00048550422,0.004739354],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989197,0.0004128611,0.00005355006,0.00020324622,0.00031405623,0.000096645905],"domain_scores_gemma":[0.9959208,0.0018336677,0.00046102662,0.0011461702,0.0004909396,0.00014748315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016488087,0.00043278866,0.00039225395,0.002325228,0.0010166619,0.0034948194,0.0017019721,0.0014792574,0.0036224958],"category_scores_gemma":[0.0055904603,0.0006638061,0.0010278049,0.0018983695,0.0026904868,0.007674898,0.0022614163,0.0012164859,0.0008660788],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008083038,0.0001120186,0.01383325,0.00014320826,0.000054759945,0.0005871057,0.0059587685,0.083949566,0.007611216,0.8302817,0.002106463,0.055281136],"study_design_scores_gemma":[0.000029349003,0.00008025856,0.0038183706,0.000059619775,0.000066393935,0.00033426654,0.001307822,0.5968795,0.0022067653,0.3767829,0.018398747,0.000035986668],"about_ca_topic_score_codex":0.0070771757,"about_ca_topic_score_gemma":0.005275085,"teacher_disagreement_score":0.0070771757,"about_ca_system_score_codex":0.0014771353,"about_ca_system_score_gemma":0.0015220932,"threshold_uncertainty_score":0.014072001},"labels":[],"label_agreement":null},{"id":"W2104953356","doi":"10.1109/wpc.2000.852493","title":"The effect of call graph construction algorithms for object-oriented programs on automatic clustering","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Call graph; Cluster analysis; Programming language; Java; Compiler; Scalability; Theoretical computer science; Algorithm; Artificial intelligence; Database","score_opus":0.015628114382333018,"score_gpt":0.26191858709244903,"score_spread":0.246290472710116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104953356","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53337276,0.0012239128,0.43870655,0.0010296742,0.00017509611,0.00044588273,0.00026284586,0.015650371,0.00913288],"genre_scores_gemma":[0.52120274,0.00042903205,0.47232294,0.0002532999,0.00005043149,0.0002333459,0.0008497695,0.002831597,0.0018267932],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99148715,0.003829543,0.00048572465,0.0012756181,0.0023354385,0.0005864668],"domain_scores_gemma":[0.84955114,0.11763362,0.0064503127,0.016773362,0.008377578,0.0012140351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007662447,0.0025104675,0.0011841496,0.0026204786,0.0022665767,0.002244292,0.0024400535,0.0016885012,0.0021165411],"category_scores_gemma":[0.0678473,0.0010765704,0.0015865709,0.0030603544,0.0021571808,0.004284439,0.0025632435,0.0023363547,0.00072215515],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019762625,0.0011690685,0.011830935,0.0005319036,0.00025371576,0.00014970357,0.00071395136,0.50005174,0.03779704,0.013597967,0.005405262,0.4265225],"study_design_scores_gemma":[0.00016821561,0.00040702493,0.0047387676,0.00005605473,0.00017168294,0.00015240847,0.00027522122,0.9280028,0.051452518,0.011911206,0.0025775873,0.000086554224],"about_ca_topic_score_codex":0.007917911,"about_ca_topic_score_gemma":0.011387715,"teacher_disagreement_score":0.007917911,"about_ca_system_score_codex":0.00254366,"about_ca_system_score_gemma":0.0026562044,"threshold_uncertainty_score":0.04052335},"labels":[],"label_agreement":null},{"id":"W2104965001","doi":"10.1109/scam.2015.7335420","title":"SimNav: Simulink navigation of model clone classes","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Computer science; Graphical user interface; Interface (matter); Context (archaeology); User interface; Model driven development; Human–computer interaction; Programming language; Software engineering; Operating system; Software; Unified Modeling Language; Biology","score_opus":0.06510927063210013,"score_gpt":0.32126370385658826,"score_spread":0.2561544332244881,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104965001","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014671765,0.00010793373,0.83494276,0.00014358517,0.00006735139,0.00013929007,0.0013211139,0.14083242,0.0077737086],"genre_scores_gemma":[0.29153788,0.0003679931,0.652693,0.00023622913,0.000032096057,0.00053497974,0.0043331706,0.029544966,0.020719728],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994006,0.00010976057,0.000039116185,0.0001415091,0.00026886453,0.00004018427],"domain_scores_gemma":[0.99769634,0.0011824259,0.00018651669,0.00045285083,0.0003860067,0.00009579728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011457567,0.0011313569,0.0005304476,0.0010336471,0.00030716124,0.0012464687,0.0016706494,0.0009276939,0.02879322],"category_scores_gemma":[0.006411158,0.0007200907,0.0006094554,0.0004138172,0.0004793847,0.0020182752,0.0017604749,0.0009453961,0.0047884905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026377991,0.0003989727,0.0106092,0.0015291655,0.00016423053,0.0013883349,0.00556728,0.07905279,0.10583071,0.067392915,0.12213674,0.60329187],"study_design_scores_gemma":[0.00046555215,0.00035508943,0.0018972865,0.0003143066,0.000073178795,0.00076496135,0.00032893723,0.50441617,0.14257267,0.014588241,0.33408332,0.0001403377],"about_ca_topic_score_codex":0.0023409375,"about_ca_topic_score_gemma":0.002523543,"teacher_disagreement_score":0.02879322,"about_ca_system_score_codex":0.00048714102,"about_ca_system_score_gemma":0.00090142223,"threshold_uncertainty_score":0.096322894},"labels":[],"label_agreement":null},{"id":"W2104994910","doi":"10.1109/tse.2012.71","title":"Trustrace: Mining Software Repositories to Improve the Accuracy of Requirement Traceability Links","year":2012,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":105,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Traceability; Computer science; Requirements traceability; Precision and recall; Source code; Software; Data mining; Software engineering; Software development; Database; Information retrieval; Programming language; Requirement","score_opus":0.020912743124051024,"score_gpt":0.27159741001598686,"score_spread":0.25068466689193586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104994910","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24205667,0.0028220494,0.70421004,0.0007849833,0.000123654,0.00045568417,0.0030681002,0.043277077,0.0032017329],"genre_scores_gemma":[0.63313824,0.0007535007,0.35141376,0.00017558479,0.000054871714,0.00027522363,0.010684615,0.0010584494,0.0024457583],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.991542,0.001645853,0.0010442545,0.0013810429,0.0040021017,0.00038475555],"domain_scores_gemma":[0.95435274,0.017463408,0.007647164,0.009763036,0.010247598,0.0005260101],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006811192,0.0018642737,0.0015505361,0.012617023,0.00094947766,0.00257971,0.0027703426,0.0013788294,0.0007195793],"category_scores_gemma":[0.058251232,0.0007106504,0.0015270152,0.006128544,0.0006426151,0.0069535575,0.0030546826,0.0014379286,0.00092555373],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005235771,0.0008655627,0.08688509,0.0014889684,0.0005211903,0.000842177,0.0017727795,0.061769076,0.021602958,0.004016467,0.013338029,0.8063741],"study_design_scores_gemma":[0.000086900276,0.0005223567,0.019618638,0.00016542485,0.00024206971,0.00087020313,0.00070098386,0.91096,0.048681796,0.008971738,0.009039544,0.00014038033],"about_ca_topic_score_codex":0.00931622,"about_ca_topic_score_gemma":0.008612623,"teacher_disagreement_score":0.012617023,"about_ca_system_score_codex":0.00092570274,"about_ca_system_score_gemma":0.0023928757,"threshold_uncertainty_score":0.03602141},"labels":[],"label_agreement":null},{"id":"W2105023644","doi":"10.1109/icsm.2001.972754","title":"A graph pattern matching approach to software architecture recovery","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Graph; Control flow graph; Theoretical computer science; Pattern matching; Source code; Software system; Software architecture; Matching (statistics); Data flow diagram; Programming language; Software; Data mining; Database","score_opus":0.019335470036792396,"score_gpt":0.22277124385389516,"score_spread":0.20343577381710276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105023644","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014083624,0.000101550344,0.99711835,0.00012072738,0.000018591976,0.00006152518,0.00003583349,0.0005778292,0.00055727153],"genre_scores_gemma":[0.026315672,0.00028879932,0.97092336,0.00009348351,0.000029332667,0.0001090709,0.00027447133,0.00014814612,0.0018176716],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973569,0.00075407774,0.00021113417,0.00047289714,0.0010517758,0.00015325395],"domain_scores_gemma":[0.99720603,0.0011281102,0.00024610755,0.0008374766,0.00050568354,0.000076591976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019275389,0.0009096846,0.0009852678,0.0041665747,0.0011880752,0.0014969533,0.002566697,0.0012867722,0.0038024632],"category_scores_gemma":[0.006206381,0.00057122554,0.0016695688,0.0043708733,0.0018037544,0.0036868807,0.0016852259,0.0019306744,0.0011735271],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001974893,0.00028645626,0.0012337232,0.0004472159,0.00015813162,0.00042510856,0.00051130867,0.07578684,0.016924208,0.17915654,0.0073600477,0.71751285],"study_design_scores_gemma":[0.00009926769,0.00023343954,0.00082213595,0.00010586569,0.00012363304,0.0011590073,0.00024664114,0.60640174,0.024754409,0.32085952,0.045115042,0.00007944438],"about_ca_topic_score_codex":0.004694641,"about_ca_topic_score_gemma":0.004769764,"teacher_disagreement_score":0.004694641,"about_ca_system_score_codex":0.0010489553,"about_ca_system_score_gemma":0.0016483772,"threshold_uncertainty_score":0.012720466},"labels":[],"label_agreement":null},{"id":"W2105182894","doi":"10.1109/icsmc.2007.4413861","title":"Benchmarking usability of early designs using predictive metrics","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Usability; Benchmarking; Computer science; Learnability; Usability inspection; Cognitive walkthrough; Usability goals; Software deployment; Component-based usability testing; Usability engineering; Heuristic evaluation; Software engineering; Human–computer interaction","score_opus":0.06879968612734985,"score_gpt":0.3205121756188769,"score_spread":0.251712489491527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105182894","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45556778,0.0010667169,0.5341111,0.00028040304,0.00007710013,0.0008669704,0.00045052145,0.0022339958,0.0053453525],"genre_scores_gemma":[0.82208526,0.00030606304,0.17522018,0.00005702041,0.000024171493,0.00066694105,0.00086452015,0.00021624217,0.00055956293],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9609491,0.0200045,0.002592068,0.002462583,0.013400788,0.00059101457],"domain_scores_gemma":[0.77235067,0.16148748,0.019371927,0.020983055,0.024532404,0.0012744957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02667765,0.0025574532,0.0014237406,0.007947372,0.0006216775,0.003681612,0.0022326568,0.0013224165,0.0011804664],"category_scores_gemma":[0.17987385,0.00069646025,0.0012315032,0.0031956986,0.00096536905,0.0054108812,0.0016788834,0.0016190544,0.0005894037],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00117048,0.0016362473,0.18067992,0.0015138753,0.0007596836,0.00024551284,0.0031123434,0.18097924,0.01598715,0.011063084,0.0020859367,0.6007665],"study_design_scores_gemma":[0.0001368791,0.0047295266,0.09044918,0.0005923675,0.00034203997,0.00036394157,0.0010089795,0.85295486,0.026631175,0.017414158,0.0049718968,0.00040503254],"about_ca_topic_score_codex":0.002395778,"about_ca_topic_score_gemma":0.0025165156,"teacher_disagreement_score":0.02667765,"about_ca_system_score_codex":0.0022529487,"about_ca_system_score_gemma":0.0013443232,"threshold_uncertainty_score":0.14108658},"labels":[],"label_agreement":null},{"id":"W2105376034","doi":"10.1145/602461.602482","title":"Growth, evolution, and structural change in open source software","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":101,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Linux kernel; Software engineering; Software; Software system; Beagle; Software maintenance; Software development; Operating system; Software construction; Software analytics","score_opus":0.03354574900947607,"score_gpt":0.2695339567565724,"score_spread":0.2359882077470963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105376034","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9429982,0.0042850263,0.030657059,0.0031562026,0.0000521076,0.00007159707,0.00028338152,0.00020603019,0.018290458],"genre_scores_gemma":[0.9907318,0.0010282127,0.006620309,0.000056539135,0.000056525816,0.00003426843,0.0002105636,0.000058681497,0.0012030342],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99614125,0.0010018732,0.0003573151,0.00060620345,0.0015638944,0.00032947652],"domain_scores_gemma":[0.948125,0.030110532,0.01257243,0.002254151,0.005357797,0.0015800062],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.004482689,0.00031603803,0.00032195303,0.0063155554,0.002221566,0.004115593,0.00094768626,0.0013226023,0.0018670061],"category_scores_gemma":[0.046849564,0.00035798055,0.00035141633,0.0072251223,0.004649771,0.01203035,0.0026586915,0.0012177008,0.00021431012],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024152332,0.00020410001,0.5813357,0.0004518718,0.00007270127,0.0018517327,0.024125349,0.01797431,0.0049083233,0.1366484,0.0021199852,0.230066],"study_design_scores_gemma":[0.00002448659,0.00028267404,0.64407355,0.00026321408,0.000075399046,0.0026617257,0.017897533,0.075136475,0.0034471445,0.22741997,0.028602215,0.00011554723],"about_ca_topic_score_codex":0.0063965674,"about_ca_topic_score_gemma":0.007200251,"teacher_disagreement_score":0.9977784,"about_ca_system_score_codex":0.0035477674,"about_ca_system_score_gemma":0.0010836192,"threshold_uncertainty_score":0.02574104},"labels":[],"label_agreement":null},{"id":"W2105492891","doi":"10.1145/2020976.2020979","title":"Managing technical debt in software development","year":2011,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Engineering management; Software quality; Debt; Software engineering; Software development; Software; Computer science; Quality (philosophy); Software project management; Engineering; Value (mathematics); Business; Software construction; Finance","score_opus":0.03005410945615254,"score_gpt":0.24478996928177618,"score_spread":0.21473585982562365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105492891","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30076033,0.018860394,0.43688083,0.084192194,0.0011017753,0.0005986654,0.00014189427,0.0012722723,0.15619169],"genre_scores_gemma":[0.91373986,0.0036783838,0.06492423,0.0027352832,0.0004671804,0.0003437516,0.00009683698,0.00028163154,0.013732866],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.977246,0.012751931,0.0017668806,0.0013770496,0.005414624,0.0014435342],"domain_scores_gemma":[0.9266963,0.03278059,0.013322833,0.008975322,0.011668889,0.0065560364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02902626,0.0005315061,0.0006374181,0.0022419505,0.005268351,0.010752572,0.0023922704,0.0031817243,0.0033710133],"category_scores_gemma":[0.10371002,0.00081871415,0.0003912926,0.0039916127,0.0049280277,0.018297782,0.014842605,0.0051270165,0.0009055604],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000295593,0.0003897993,0.021550452,0.00080899155,0.00007045122,0.0014647582,0.022953818,0.013858378,0.0028256,0.38968062,0.026969016,0.5191326],"study_design_scores_gemma":[0.00017256112,0.00044894384,0.010814388,0.0012548045,0.000085961685,0.0018807382,0.011286383,0.041599937,0.0024354672,0.75579023,0.17407696,0.00015364101],"about_ca_topic_score_codex":0.0026373887,"about_ca_topic_score_gemma":0.0026702262,"teacher_disagreement_score":0.02902626,"about_ca_system_score_codex":0.0047512525,"about_ca_system_score_gemma":0.006761673,"threshold_uncertainty_score":0.15350735},"labels":[],"label_agreement":null},{"id":"W2105638048","doi":"10.1145/1328279.1328294","title":"Informing Eclipse API production and consumption","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Eclipse; Application programming interface; Simple (philosophy); Software engineering; Programming language; World Wide Web","score_opus":0.023883946826838278,"score_gpt":0.2828578363998166,"score_spread":0.25897388957297834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105638048","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14732501,0.0006367402,0.6168036,0.0035344458,0.0005284627,0.0009786108,0.009214127,0.10027832,0.120700724],"genre_scores_gemma":[0.43373775,0.0004875636,0.49534068,0.0004093979,0.00012506965,0.00093196577,0.010814975,0.02305462,0.03509802],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99303985,0.002458527,0.0006576122,0.0006991713,0.0026852824,0.0004595659],"domain_scores_gemma":[0.9495749,0.027313348,0.002623243,0.011759256,0.008074569,0.00065476703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014302792,0.00097317924,0.00064504927,0.003148295,0.00090166583,0.0033667951,0.0010293856,0.0012108502,0.010423285],"category_scores_gemma":[0.0736879,0.0014148011,0.0004213995,0.0014006296,0.00052980863,0.0059320815,0.0018480258,0.0016422325,0.004558712],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016972073,0.00035215044,0.04735785,0.0007326972,0.000061920255,0.00086919713,0.006723079,0.019912232,0.035306074,0.07057021,0.084973544,0.7314438],"study_design_scores_gemma":[0.00031944038,0.0003597255,0.031213716,0.0005455355,0.00011525375,0.0007533255,0.0025983984,0.25230458,0.10167073,0.0440253,0.5657734,0.00032066726],"about_ca_topic_score_codex":0.007973709,"about_ca_topic_score_gemma":0.013832919,"teacher_disagreement_score":0.014302792,"about_ca_system_score_codex":0.0020644786,"about_ca_system_score_gemma":0.0035570238,"threshold_uncertainty_score":0.075641274},"labels":[],"label_agreement":null},{"id":"W2105672266","doi":"10.1145/1852786.1852792","title":"Understanding the impact of code and process metrics on post-release defects","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":94,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Merge (version control); Software quality; Software metric; Data mining; Eclipse; Source code; Process (computing); Software; Data science; Machine learning; Software development; Information retrieval; Programming language","score_opus":0.056479178751535304,"score_gpt":0.3251365449346347,"score_spread":0.2686573661830994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105672266","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8797784,0.00041540404,0.11622188,0.0006058105,0.00002368975,0.00007674016,0.0005047511,0.00079384336,0.0015794245],"genre_scores_gemma":[0.9897559,0.00014504026,0.009215545,0.000025981504,0.000012588132,0.00002662972,0.00042518007,0.000062985644,0.00033015307],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9954495,0.0020686951,0.00023197613,0.0007966507,0.0011278728,0.00032520993],"domain_scores_gemma":[0.8319849,0.1372962,0.01664263,0.006741744,0.0063332524,0.0010012569],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007652773,0.0015597058,0.00069335004,0.004579549,0.00030364515,0.0016982388,0.00088311825,0.0016409579,0.0013092833],"category_scores_gemma":[0.08255118,0.0006946943,0.0012816839,0.0021912563,0.000690972,0.0036982428,0.00097016874,0.001715075,0.00042908493],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046367297,0.0005358935,0.59691185,0.0002180289,0.00049231696,0.0004979184,0.00079370773,0.2811986,0.0058600428,0.004282796,0.000984517,0.10776068],"study_design_scores_gemma":[0.000014261959,0.0003035866,0.17521197,0.000029549776,0.000108422675,0.00016302361,0.00014566243,0.8164179,0.002653258,0.0044613895,0.00045032313,0.00004064443],"about_ca_topic_score_codex":0.008818725,"about_ca_topic_score_gemma":0.010815985,"teacher_disagreement_score":0.008818725,"about_ca_system_score_codex":0.00093084096,"about_ca_system_score_gemma":0.0010940171,"threshold_uncertainty_score":0.04047221},"labels":[],"label_agreement":null},{"id":"W2105756730","doi":"10.1109/wse.2008.4655403","title":"An approach for estimating code changes in e-commerce applications","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Business process modeling; Artifact-centric business process model; Business process; Business process management; Business process discovery; Business rule; Business Process Model and Notation; Metric (unit); Source code; Tracing; Code (set theory); Process (computing); Process management; Software engineering; Work in process; Business; Programming language; Set (abstract data type); Marketing","score_opus":0.06274594035159221,"score_gpt":0.322540788508306,"score_spread":0.2597948481567138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105756730","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16810955,0.0006623407,0.82222027,0.00022604897,0.00006193872,0.00072679477,0.0009956197,0.0040373905,0.0029599953],"genre_scores_gemma":[0.38219056,0.00024592978,0.6146229,0.000043653745,0.00002740612,0.00037046566,0.0013329077,0.0002533411,0.0009129207],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943733,0.00091437146,0.0005513314,0.00083651324,0.0031442072,0.00018024101],"domain_scores_gemma":[0.9740263,0.010832008,0.0046379045,0.0022438865,0.007849393,0.00041054364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029002826,0.0012553098,0.00069010694,0.011371118,0.0007214524,0.0016750268,0.001338771,0.0014112325,0.0007024803],"category_scores_gemma":[0.03797218,0.0007082749,0.0009318459,0.0060300515,0.0005689857,0.0027603537,0.0012037376,0.0014840453,0.00042440754],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025958367,0.00057783263,0.13345164,0.00054810295,0.0003078503,0.0003763453,0.0010943305,0.17168327,0.035121486,0.0076894537,0.0020353815,0.64685476],"study_design_scores_gemma":[0.00003491211,0.00044491806,0.06359833,0.00006220694,0.00012377111,0.00047042978,0.0003048719,0.8853289,0.036585703,0.007106053,0.005847499,0.000092435446],"about_ca_topic_score_codex":0.009843887,"about_ca_topic_score_gemma":0.008973891,"teacher_disagreement_score":0.011371118,"about_ca_system_score_codex":0.0015544158,"about_ca_system_score_gemma":0.0014765862,"threshold_uncertainty_score":0.019573152},"labels":[],"label_agreement":null},{"id":"W2105827483","doi":"10.1109/vlhcc.2008.4639095","title":"How can diagramming tools help support programming activities?","year":2008,"lang":"en","type":"article","venue":"Proceedings/Proceedings -- IEEE Symposium on Visual Languages and Human-Centric Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Diagrammatic reasoning; Computer science; Exploratory research; Context (archaeology); Human–computer interaction; Software engineering; Knowledge management; Programming language","score_opus":0.022532426139985545,"score_gpt":0.2910056960305173,"score_spread":0.2684732698905318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105827483","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61992204,0.0050279377,0.30809054,0.022520049,0.0003513735,0.00040418975,0.00023972824,0.0070124734,0.03643157],"genre_scores_gemma":[0.781645,0.0026230533,0.21096122,0.0010143527,0.00010900482,0.00021669928,0.00022189252,0.0004755726,0.0027331137],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99137115,0.006125164,0.00028740327,0.000745206,0.0010168836,0.00045427872],"domain_scores_gemma":[0.90515274,0.080684744,0.004194179,0.0042883237,0.0037949388,0.0018850252],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0103405155,0.0011593822,0.0004912246,0.0016888726,0.00083651347,0.0056865728,0.0020136866,0.0026497683,0.002502312],"category_scores_gemma":[0.10700979,0.0006523818,0.00040281436,0.0014093581,0.0013322332,0.011050055,0.0016786372,0.0014126934,0.0017499594],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047324528,0.0012971763,0.021755543,0.0028272986,0.000099780445,0.00071361306,0.037700605,0.002834937,0.020648275,0.014593652,0.009615473,0.88744044],"study_design_scores_gemma":[0.0008914875,0.004076198,0.05199944,0.006071847,0.0006407726,0.006891787,0.10102496,0.06117143,0.061859384,0.28096032,0.42372304,0.0006893433],"about_ca_topic_score_codex":0.00062454795,"about_ca_topic_score_gemma":0.0011146447,"teacher_disagreement_score":0.0103405155,"about_ca_system_score_codex":0.00053277216,"about_ca_system_score_gemma":0.0011816842,"threshold_uncertainty_score":0.054686487},"labels":[],"label_agreement":null},{"id":"W2105889333","doi":"10.1109/icpc.2006.18","title":"Digging the Development Dust for Refactorings","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Code refactoring; Computer science; Context (archaeology); Software development; Software engineering; Identifier; Software development process; Software; Software evolution; Process (computing); Source code; Data science; Software construction; Programming language","score_opus":0.024502564607944446,"score_gpt":0.26121014348600324,"score_spread":0.2367075788780588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105889333","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8904519,0.003101866,0.08242809,0.009167091,0.00010975978,0.00018591699,0.0005636936,0.00076278934,0.013228828],"genre_scores_gemma":[0.9178262,0.0016833023,0.0735272,0.0005465023,0.000068205474,0.00009941476,0.00075810787,0.00035912276,0.0051319776],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98710465,0.003126919,0.0009710773,0.0017394341,0.0064737867,0.0005841558],"domain_scores_gemma":[0.7928243,0.13301794,0.033261992,0.021243813,0.01706987,0.0025820304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013893388,0.00073800015,0.0008454762,0.010226269,0.0027359452,0.0069286996,0.0021538055,0.0024428058,0.0037807222],"category_scores_gemma":[0.13335219,0.0015767824,0.00079842366,0.005230439,0.0020955328,0.01569145,0.003721272,0.0028563302,0.0008740782],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000348058,0.00034831665,0.25362703,0.0007946825,0.00024770308,0.0037626673,0.018105373,0.0068826396,0.0065227533,0.028428204,0.004130878,0.67680174],"study_design_scores_gemma":[0.00015001566,0.001878707,0.46757522,0.003987532,0.0006633605,0.010628973,0.037512675,0.07869203,0.029918175,0.22306503,0.14518102,0.00074734224],"about_ca_topic_score_codex":0.0035263496,"about_ca_topic_score_gemma":0.010292458,"teacher_disagreement_score":0.013893388,"about_ca_system_score_codex":0.0032763674,"about_ca_system_score_gemma":0.004160487,"threshold_uncertainty_score":0.073476076},"labels":[],"label_agreement":null},{"id":"W2105899414","doi":"10.1007/s10515-008-0028-6","title":"Requirements model generation to support requirements elicitation: the Secure Tropos experience","year":2008,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Requirements elicitation; Requirements engineering; Computer science; Requirements analysis; Requirements management; Context (archaeology); Process (computing); Software engineering; Requirement; Systems engineering; Engineering","score_opus":0.06163986956967274,"score_gpt":0.3091074000294848,"score_spread":0.24746753045981204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105899414","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12617803,0.00044053677,0.8428666,0.0024198578,0.000087380344,0.00042227266,0.0004165047,0.010821784,0.016347088],"genre_scores_gemma":[0.37291318,0.000692484,0.61769843,0.00044043525,0.000041028456,0.0002646239,0.0010358721,0.0018359757,0.0050778943],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99504155,0.0032794133,0.00020136956,0.00030540655,0.0010387948,0.00013347136],"domain_scores_gemma":[0.9753067,0.017948747,0.0008630948,0.0040934146,0.001324654,0.0004634422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008043161,0.0007707202,0.0004406258,0.0007164153,0.0006110636,0.0016844759,0.0018124906,0.0012627533,0.0046519316],"category_scores_gemma":[0.017942702,0.0006974134,0.00074767234,0.00051380834,0.0012672475,0.004012785,0.0023455312,0.0022207417,0.0016045728],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025006344,0.0014223097,0.0049029286,0.0011572799,0.00025867645,0.0020807104,0.013877714,0.12074629,0.09631364,0.13314477,0.031322535,0.5922726],"study_design_scores_gemma":[0.00081912597,0.0008733537,0.001216081,0.00031762707,0.0001346333,0.0013790592,0.0027056886,0.7567119,0.04777127,0.088170156,0.0997748,0.00012628453],"about_ca_topic_score_codex":0.001605693,"about_ca_topic_score_gemma":0.0019220219,"teacher_disagreement_score":0.008043161,"about_ca_system_score_codex":0.00067966926,"about_ca_system_score_gemma":0.0016426734,"threshold_uncertainty_score":0.042536795},"labels":[],"label_agreement":null},{"id":"W2105950237","doi":"10.1145/2791060.2791107","title":"Maintaining feature traceability with embedded annotations","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Traceability; Computer science; Feature (linguistics); Reuse; Code reuse; Software; Annotation; Software maintenance; Feature model; Code (set theory); Software development; Source lines of code; Software product line; Software engineering; Data mining; Artificial intelligence; Programming language; Engineering","score_opus":0.028747973279675463,"score_gpt":0.2842272505381935,"score_spread":0.255479277258518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105950237","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7808076,0.00012096106,0.21037182,0.000316027,0.00004849235,0.00018038807,0.00032938336,0.0049711578,0.0028542138],"genre_scores_gemma":[0.8568539,0.000081188744,0.13959353,0.000058690082,0.000010782016,0.00016414237,0.0006659639,0.00064737,0.001924461],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99650323,0.0011395036,0.00029777674,0.000508627,0.0012433305,0.00030751998],"domain_scores_gemma":[0.9541206,0.023832273,0.0047334344,0.0126054445,0.0040289112,0.00067939656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039069587,0.0010159744,0.00045508178,0.0010688474,0.0006490227,0.0014959429,0.002010261,0.0013864256,0.0016750116],"category_scores_gemma":[0.040334344,0.00072993763,0.0007001357,0.0009637484,0.0011705094,0.0032450517,0.0020187122,0.0014784749,0.00043451533],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013515223,0.0013826942,0.0898635,0.0006886541,0.00016670389,0.0012733919,0.0025369565,0.52923226,0.13325569,0.011195728,0.0023906003,0.22666228],"study_design_scores_gemma":[0.00012362754,0.0008867627,0.017736537,0.00008434445,0.00021425211,0.00044583651,0.0003893591,0.90280503,0.057115592,0.012546546,0.0075482037,0.000103872175],"about_ca_topic_score_codex":0.0071146786,"about_ca_topic_score_gemma":0.006918798,"teacher_disagreement_score":0.0071146786,"about_ca_system_score_codex":0.001086762,"about_ca_system_score_gemma":0.0023571386,"threshold_uncertainty_score":0.020662248},"labels":[],"label_agreement":null},{"id":"W2106013472","doi":"10.1109/wcre.2002.1173068","title":"Java quality assurance by detecting code smells","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":380,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Code smell; Code refactoring; Computer science; Static program analysis; Software quality; Software engineering; KPI-driven code analysis; Code review; Software visualization; Code (set theory); Program comprehension; Software inspection; Source code; Software; Java; Programming language; Software development; Software system; Software construction","score_opus":0.029908899639837307,"score_gpt":0.3101992349853636,"score_spread":0.28029033534552633,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106013472","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46330482,0.0012690162,0.44697192,0.0013847252,0.0001351081,0.00045988586,0.0007098764,0.07816157,0.0076031517],"genre_scores_gemma":[0.74303836,0.00051512284,0.25127083,0.00018225395,0.00005100859,0.00011516778,0.00077424286,0.0021487374,0.001904196],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955285,0.0011623697,0.0003864171,0.0006435195,0.0020937836,0.00018543765],"domain_scores_gemma":[0.94742906,0.021285998,0.012639279,0.0067578414,0.010813366,0.0010744698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046594203,0.0007628106,0.00064578187,0.0033726864,0.00048576016,0.0024119434,0.00087694003,0.0008523011,0.0009009492],"category_scores_gemma":[0.040155396,0.00056813034,0.0003659403,0.0014049688,0.00060370273,0.0023769,0.0014364539,0.001310785,0.000583895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010419212,0.00044099422,0.12146056,0.0012215571,0.00014590178,0.0010593681,0.0052539096,0.010853073,0.16557571,0.0033614677,0.011277748,0.6783078],"study_design_scores_gemma":[0.00026333993,0.0014296756,0.25766924,0.00082859397,0.00037581357,0.003967044,0.0017311462,0.40878424,0.26948535,0.011563861,0.04334979,0.00055194256],"about_ca_topic_score_codex":0.0017826782,"about_ca_topic_score_gemma":0.002327667,"teacher_disagreement_score":0.0046594203,"about_ca_system_score_codex":0.00051395356,"about_ca_system_score_gemma":0.0009008547,"threshold_uncertainty_score":0.024641633},"labels":[],"label_agreement":null},{"id":"W2106029414","doi":"10.1109/icsm.2004.1357808","title":"Understanding phases and styles of object-oriented systems' evolution","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Software evolution; Computer science; Software; Software system; Structural pattern; Matching (statistics); Software development; Software maintenance; Software design; Programming language; Mathematics; Software construction","score_opus":0.05345162739433197,"score_gpt":0.26893454680236273,"score_spread":0.21548291940803077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106029414","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77535343,0.00048422857,0.21248947,0.00044212438,0.000015584455,0.00029872567,0.00021946113,0.00027830782,0.010418669],"genre_scores_gemma":[0.8732611,0.00035020694,0.12440699,0.000044899305,0.000011883434,0.00013642725,0.00022803816,0.000057363646,0.0015031549],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9985216,0.0005699942,0.000119485485,0.00015176374,0.00055362453,0.000083517625],"domain_scores_gemma":[0.989814,0.0052632075,0.0021619657,0.0011128028,0.0013850295,0.00026297593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024672141,0.0002943374,0.00017286191,0.0023579393,0.00056286604,0.0023186211,0.0004425835,0.00063840335,0.00080386223],"category_scores_gemma":[0.0146556515,0.0003507603,0.00033291295,0.0016625673,0.000981113,0.0036101632,0.00070256396,0.0006350293,0.00020161665],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031771034,0.00024331962,0.25115377,0.0003989876,0.00007372039,0.0005635989,0.020209648,0.025072653,0.028686976,0.13768938,0.0010760304,0.5345141],"study_design_scores_gemma":[0.00007021199,0.0004959394,0.26175016,0.0003013549,0.00011006989,0.0028053117,0.015714275,0.22864327,0.026073255,0.4298446,0.03402984,0.00016174157],"about_ca_topic_score_codex":0.0012056924,"about_ca_topic_score_gemma":0.002250608,"teacher_disagreement_score":0.0024672141,"about_ca_system_score_codex":0.00067690905,"about_ca_system_score_gemma":0.0005959575,"threshold_uncertainty_score":0.013048053},"labels":[],"label_agreement":null},{"id":"W2106095006","doi":"10.1142/s021819401000492x","title":"DYNAMIC KNOWLEDGE EXTRACTION FROM SOFTWARE SYSTEMS USING SEQUENTIAL PATTERN MINING","year":2010,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Software system; Software construction; Source code; Software sizing; Software; Cohesion (chemistry); Software visualization; Software framework; Software development; Software engineering; Static program analysis; Software metric; Component-based software engineering; Data mining; Programming language","score_opus":0.014690367136132271,"score_gpt":0.2840261469821888,"score_spread":0.26933577984605656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106095006","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.079453625,0.0007260108,0.91094875,0.00039151002,0.000037726506,0.00051846844,0.0023228035,0.003486755,0.0021143327],"genre_scores_gemma":[0.26218376,0.0007872328,0.7285532,0.00007213835,0.000026761254,0.00041824215,0.0067827287,0.0001558148,0.0010200578],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998007,0.00033776593,0.00027421553,0.00044572543,0.0008183917,0.00011686635],"domain_scores_gemma":[0.9954045,0.0023421135,0.00059588504,0.000719123,0.00083564,0.00010270048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011505092,0.0011967021,0.0009731932,0.009567321,0.0006655777,0.001355924,0.0013246905,0.0007259418,0.0010351894],"category_scores_gemma":[0.0068729604,0.0005159712,0.0014452332,0.0063091367,0.00060139847,0.0019725347,0.0011730632,0.0007231902,0.0007954524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021091053,0.0003350909,0.017831111,0.0013674519,0.00026465082,0.0017180288,0.0008322015,0.048994087,0.02769135,0.0072784955,0.0039050637,0.88957155],"study_design_scores_gemma":[0.000095016585,0.00041590288,0.014359774,0.0003001567,0.0003608637,0.0025412824,0.0011449766,0.8308871,0.056447554,0.06993409,0.023394743,0.00011856341],"about_ca_topic_score_codex":0.004059378,"about_ca_topic_score_gemma":0.005768519,"teacher_disagreement_score":0.009567321,"about_ca_system_score_codex":0.00074168044,"about_ca_system_score_gemma":0.0019794567,"threshold_uncertainty_score":0.008071482},"labels":[],"label_agreement":null},{"id":"W2106112694","doi":"10.1145/1081706.1081744","title":"Strathcona example recommendation tool","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Calgary","funders":"","keywords":"Computer science","score_opus":0.03773330956869129,"score_gpt":0.28215116805406004,"score_spread":0.24441785848536876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106112694","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027194262,0.002341682,0.4642097,0.005295334,0.0006463911,0.002196013,0.09189982,0.23907153,0.16714527],"genre_scores_gemma":[0.07741056,0.0017638796,0.72312933,0.0015693747,0.00016631598,0.0020087333,0.088626094,0.005907825,0.099417925],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991998,0.00014425322,0.00008022006,0.00011626433,0.0004149981,0.00004456347],"domain_scores_gemma":[0.9967841,0.0017113894,0.0001104057,0.00037071834,0.00089017133,0.0001332888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000680893,0.00091702247,0.0007547346,0.0037068883,0.0007435902,0.0012823817,0.001357918,0.0012664028,0.06739416],"category_scores_gemma":[0.0058594192,0.0004727589,0.0005530822,0.0029852903,0.00017560871,0.0012925311,0.0008665559,0.0006791733,0.021562994],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048862206,0.00015533983,0.002626129,0.0012251702,0.000062277926,0.00076810515,0.00029910926,0.0028051983,0.004557424,0.009329417,0.5639219,0.41376132],"study_design_scores_gemma":[0.00034846063,0.00014160226,0.0037826537,0.00028097822,0.00007781854,0.0009848312,0.00022746951,0.05440169,0.00987889,0.0071805646,0.92258424,0.0001107912],"about_ca_topic_score_codex":0.01253935,"about_ca_topic_score_gemma":0.030526081,"teacher_disagreement_score":0.06739416,"about_ca_system_score_codex":0.00056635996,"about_ca_system_score_gemma":0.00093309773,"threshold_uncertainty_score":0.22545594},"labels":[],"label_agreement":null},{"id":"W2106419796","doi":"10.5555/504800.504802","title":"Evolution of the writer's role","year":2000,"lang":"en","type":"article","venue":"International Professional Communication Conference","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Computer science; Usability; Product (mathematics); World Wide Web; Human–computer interaction; Multimedia; Software engineering","score_opus":0.022159170773306747,"score_gpt":0.30572979170188685,"score_spread":0.2835706209285801,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106419796","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5400111,0.0038552382,0.034955196,0.03429682,0.004855418,0.00024994885,0.00018095979,0.0012442091,0.380351],"genre_scores_gemma":[0.85689193,0.00090932625,0.009856981,0.002838962,0.00072025316,0.000120613484,0.00012853609,0.00070821174,0.12782519],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97629255,0.010187947,0.0010126125,0.0045264787,0.0054738047,0.0025065707],"domain_scores_gemma":[0.88554853,0.03311158,0.008395106,0.0074745407,0.031287175,0.034183085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017119415,0.00069015624,0.0007297356,0.0041588126,0.013593942,0.017006222,0.0030982797,0.0034241925,0.01601134],"category_scores_gemma":[0.060733527,0.001000059,0.0005493977,0.001492646,0.0066814576,0.009205111,0.007615867,0.0048617595,0.008880487],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011422414,0.0018878744,0.095619045,0.0005598379,0.00011380938,0.004613244,0.26445654,0.00082035793,0.013517128,0.16863985,0.059538573,0.38909152],"study_design_scores_gemma":[0.00024639486,0.00078532356,0.03317771,0.00077247946,0.00011985237,0.0058266562,0.10443691,0.0033207678,0.005656949,0.018174324,0.8272209,0.00026184082],"about_ca_topic_score_codex":0.0047228765,"about_ca_topic_score_gemma":0.004336331,"teacher_disagreement_score":0.017119415,"about_ca_system_score_codex":0.0065393117,"about_ca_system_score_gemma":0.011397587,"threshold_uncertainty_score":0.09053725},"labels":[],"label_agreement":null},{"id":"W2106571622","doi":"10.1109/wcre.2009.22","title":"SQUAD: Software Quality Understanding through the Analysis of Design","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Software quality; Computer science; Quality (philosophy); Code smell; Software engineering; Software quality control; Object-oriented design; Software; Software design pattern; Software design; Software quality analyst; Software development; Programming language","score_opus":0.15848084052111594,"score_gpt":0.359248210995698,"score_spread":0.20076737047458204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106571622","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02278114,0.00052312063,0.97006273,0.0007403479,0.00002274791,0.0001174918,0.00036266554,0.002022024,0.0033677784],"genre_scores_gemma":[0.22903885,0.00077693956,0.7666716,0.00015848622,0.000037374277,0.0003105496,0.0013023021,0.0004734612,0.001230495],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99686205,0.00095807813,0.00022504985,0.00034878514,0.0014986777,0.000107422165],"domain_scores_gemma":[0.9912952,0.004021211,0.0014630955,0.0017330678,0.0013584609,0.00012894926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004258528,0.0016414392,0.00088289817,0.0050888485,0.0007536562,0.0034683894,0.0016174458,0.0009865913,0.0026729477],"category_scores_gemma":[0.016915707,0.0008539309,0.0018024843,0.0023188132,0.00208339,0.007170911,0.0027409056,0.0018513923,0.00063206337],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013977522,0.00029300418,0.029772276,0.0017009077,0.0002920295,0.0004518553,0.005500419,0.17039081,0.011497945,0.3375003,0.010913741,0.43154705],"study_design_scores_gemma":[0.000049462877,0.0001536826,0.00480842,0.00040036804,0.0001255417,0.00031908587,0.0012458381,0.6764103,0.0057236454,0.28475642,0.025936373,0.00007085447],"about_ca_topic_score_codex":0.004091319,"about_ca_topic_score_gemma":0.0033988922,"teacher_disagreement_score":0.0050888485,"about_ca_system_score_codex":0.0019242935,"about_ca_system_score_gemma":0.0024488985,"threshold_uncertainty_score":0.022521496},"labels":[],"label_agreement":null},{"id":"W2106628621","doi":"10.1109/iciecs.2009.5364521","title":"Predicting Co-Changed Software Entities in the Context of Software Evolution","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Waterloo","keywords":"Computer science; Data mining; Software; Context (archaeology); Software evolution; Tracing; Database transaction; Matching (statistics); Software development; Software engineering; Software construction; Database; Programming language","score_opus":0.019734774778992666,"score_gpt":0.26657365557468293,"score_spread":0.24683888079569027,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106628621","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7910335,0.0012524467,0.20317149,0.00030201385,0.000049193535,0.00014149814,0.0011132058,0.0014596366,0.0014769973],"genre_scores_gemma":[0.89952314,0.00043744926,0.09756211,0.000030622068,0.00002999309,0.000035464745,0.0017376038,0.000038953596,0.0006046791],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835396,0.0003037013,0.0001967339,0.00043824862,0.0006111565,0.00009621276],"domain_scores_gemma":[0.98930115,0.0057247644,0.001761227,0.0010778297,0.0018921248,0.0002429308],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013781926,0.00056134525,0.000618393,0.0058049313,0.0004941421,0.0010064082,0.0007188562,0.001051653,0.00050899223],"category_scores_gemma":[0.012507314,0.00032048827,0.0006086859,0.0034149238,0.00022998599,0.002272574,0.00074235507,0.0006377137,0.00032737854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048137913,0.00032876033,0.49754897,0.00032470663,0.000296859,0.0017589377,0.0007541625,0.028557563,0.02038915,0.0011592848,0.0019523533,0.44644797],"study_design_scores_gemma":[0.000026504767,0.0002950695,0.23396228,0.000048278776,0.00029252766,0.002628454,0.00062425394,0.7253645,0.028967388,0.003364278,0.0043627894,0.00006369433],"about_ca_topic_score_codex":0.0038427736,"about_ca_topic_score_gemma":0.0072338833,"teacher_disagreement_score":0.0058049313,"about_ca_system_score_codex":0.00034739848,"about_ca_system_score_gemma":0.00039264152,"threshold_uncertainty_score":0.007640779},"labels":[],"label_agreement":null},{"id":"W2106648531","doi":"10.1109/tr.2014.2366274","title":"Identifying Recurring Faulty Functions in Field Traces of a Large Industrial Software System","year":2014,"lang":"en","type":"article","venue":"IEEE Transactions on Reliability","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Western University","funders":"","keywords":"Computer science; Field (mathematics); Software; Software system; Source lines of code; Software bug; Code (set theory); Software quality; Debugging; Reliability engineering; Software engineering; Software development; Programming language; Engineering","score_opus":0.027299124616901146,"score_gpt":0.2791552777762735,"score_spread":0.2518561531593724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106648531","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93075436,0.00035144656,0.06438009,0.00016644168,0.000030902014,0.00005462664,0.0006682807,0.0031655328,0.00042830617],"genre_scores_gemma":[0.9755563,0.00008053209,0.02303938,0.000019849966,0.000009586557,0.00001567214,0.00094050827,0.000049210164,0.00028898293],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9995074,0.00010540764,0.000044733777,0.00014318699,0.0001495274,0.00004979769],"domain_scores_gemma":[0.99478394,0.0028586697,0.00073271187,0.00067119865,0.000746023,0.00020739205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00081402995,0.0007143639,0.00032822796,0.0019383747,0.00031652112,0.0003947647,0.00072555593,0.00059260655,0.00032498845],"category_scores_gemma":[0.006236322,0.00026972257,0.0002906027,0.0008678094,0.00035407126,0.00069538585,0.000394428,0.0007197662,0.0002107102],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005470055,0.0008357547,0.24824314,0.0004161543,0.00016950264,0.00180915,0.0015394854,0.3023302,0.040123425,0.0008399768,0.0046116896,0.3985346],"study_design_scores_gemma":[0.000021681079,0.0002694583,0.040620886,0.00003207796,0.000030110614,0.0003705603,0.00022761925,0.9371763,0.018197523,0.0016326157,0.001392205,0.0000289503],"about_ca_topic_score_codex":0.009287579,"about_ca_topic_score_gemma":0.014306242,"teacher_disagreement_score":0.009287579,"about_ca_system_score_codex":0.000519623,"about_ca_system_score_gemma":0.0005721061,"threshold_uncertainty_score":0.018467009},"labels":[],"label_agreement":null},{"id":"W2106743089","doi":"10.1007/s11219-004-5259-6","title":"The Evolution Path for Industrial Software Quality Evaluation Methods Applying ISO/IEC 9126:2001 Quality Model: Example of MITRE?s SQAE Method","year":2005,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software quality; Reliability engineering; Quality (philosophy); Software; Computer science; Verification and validation; Software quality control; Engineering; Systems engineering; Software development; Operations management; Operating system","score_opus":0.25675761791988566,"score_gpt":0.4826588867655611,"score_spread":0.22590126884567546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106743089","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03651257,0.0005780924,0.930137,0.0020505197,0.00007654877,0.00024824403,0.00016125702,0.0007027837,0.029532928],"genre_scores_gemma":[0.2096852,0.00041611484,0.7802092,0.00016772283,0.0000141640285,0.00024829514,0.0002396073,0.00024056114,0.0087791355],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9943295,0.0022710068,0.00032481272,0.00046984226,0.002416732,0.0001881188],"domain_scores_gemma":[0.98516315,0.00565581,0.0010679147,0.0017447666,0.006124237,0.00024410333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067314296,0.00045289315,0.00038995798,0.0024058507,0.0013458194,0.002155534,0.0010179252,0.001560072,0.00541616],"category_scores_gemma":[0.027442312,0.00052004267,0.00061481027,0.002222874,0.0014014115,0.0033685523,0.0015602629,0.0018842726,0.0015586651],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012891451,0.00026084672,0.012689386,0.00039212586,0.00005391497,0.00019715523,0.0024980567,0.021247815,0.007764296,0.41019228,0.005371984,0.5392033],"study_design_scores_gemma":[0.0001134366,0.0006764407,0.015134714,0.0007646364,0.00011747424,0.0009927524,0.0014528729,0.5889753,0.017522154,0.28398144,0.09012905,0.00013974887],"about_ca_topic_score_codex":0.0055378242,"about_ca_topic_score_gemma":0.006301778,"teacher_disagreement_score":0.0067314296,"about_ca_system_score_codex":0.0022594472,"about_ca_system_score_gemma":0.0026703316,"threshold_uncertainty_score":0.03559965},"labels":[],"label_agreement":null},{"id":"W2107024044","doi":"10.1109/tse.2006.38","title":"On the value of static analysis for fault detection in software","year":2006,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":281,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nortel (Canada)","funders":"North Carolina State University; National Science Foundation","keywords":"Static analysis; Computer science; Fault detection and isolation; Software; Software quality; Reliability engineering; Software bug; Programmer; Software reliability testing; Fault (geology); Static program analysis; Data mining; Software development; Embedded system; Operating system; Artificial intelligence; Programming language; Engineering","score_opus":0.011465499236457887,"score_gpt":0.23584611802258723,"score_spread":0.22438061878612933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107024044","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44364768,0.0064586787,0.53238446,0.002100173,0.00012990429,0.00017407352,0.00024011816,0.0024782375,0.0123866685],"genre_scores_gemma":[0.92213583,0.0011519614,0.07539402,0.00015196965,0.00009865635,0.00005144287,0.000113587106,0.00013386167,0.00076870475],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98981655,0.004996632,0.00038934598,0.0005871852,0.00392152,0.00028864158],"domain_scores_gemma":[0.89125115,0.091361634,0.0039415,0.006717937,0.006354789,0.0003730584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006543355,0.0011200513,0.0007380307,0.00580374,0.0007779919,0.002279223,0.00087680964,0.000882963,0.0014516666],"category_scores_gemma":[0.042260434,0.00047726117,0.0007592982,0.0029583767,0.0034785962,0.0049676676,0.0007291718,0.0009961132,0.00054153294],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083894783,0.00035182276,0.078422844,0.0004414164,0.0003010699,0.00045129014,0.0006417152,0.17200801,0.019094637,0.024888566,0.0015298467,0.7010298],"study_design_scores_gemma":[0.0000784717,0.0016310149,0.055131998,0.00038462356,0.00039311356,0.0014348156,0.0005844062,0.7966369,0.037855394,0.09934291,0.0061585098,0.00036782015],"about_ca_topic_score_codex":0.002920281,"about_ca_topic_score_gemma":0.0025826269,"teacher_disagreement_score":0.006543355,"about_ca_system_score_codex":0.001047457,"about_ca_system_score_gemma":0.0011623607,"threshold_uncertainty_score":0.034604967},"labels":[],"label_agreement":null},{"id":"W2107079484","doi":"10.1109/ase.2003.1240290","title":"Analysis of inconsistency in graph-based viewpoints: a category-theoretical approach","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Viewpoints; Computer science; Merge (version control); Fuzzy logic; Theoretical computer science; Graph; Vagueness; Artificial intelligence; Data mining; Information retrieval","score_opus":0.013058292515827003,"score_gpt":0.2547652757188619,"score_spread":0.24170698320303494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107079484","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018174203,0.00015785247,0.97893184,0.000390664,0.000016216867,0.00007358162,0.000042536052,0.00008608431,0.0021270511],"genre_scores_gemma":[0.3568425,0.0002875766,0.64115125,0.00015481013,0.000054268632,0.00032932858,0.0002577198,0.00007184031,0.00085067534],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9854757,0.0067911292,0.0011011969,0.0014656457,0.004510287,0.00065593503],"domain_scores_gemma":[0.9637552,0.02328109,0.0026892761,0.00477462,0.004667781,0.00083202467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015823366,0.000870649,0.001199261,0.011333527,0.0029775002,0.0055011134,0.0035894508,0.0026301988,0.0020421136],"category_scores_gemma":[0.031244116,0.0009664771,0.0026962275,0.00567186,0.009282166,0.011021655,0.00585925,0.0027655226,0.0001910573],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004537075,0.000031990847,0.002073988,0.00017203535,0.00008715125,0.00033622805,0.0027544533,0.013986908,0.0018121554,0.9519824,0.00031700617,0.026400315],"study_design_scores_gemma":[0.00001817545,0.00005443665,0.0008162885,0.000107211075,0.00009052114,0.00035699183,0.0015271817,0.06713513,0.0020737755,0.92221624,0.005548105,0.00005601354],"about_ca_topic_score_codex":0.0040446417,"about_ca_topic_score_gemma":0.0023491313,"teacher_disagreement_score":0.015823366,"about_ca_system_score_codex":0.003919579,"about_ca_system_score_gemma":0.0022917467,"threshold_uncertainty_score":0.083682954},"labels":[],"label_agreement":null},{"id":"W2107142491","doi":"10.1145/1718918.1718973","title":"Information needs in bug reports","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":203,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software bug; Eclipse; Computer science; Software engineering; Sample (material); Tracking (education); Software; Data science; World Wide Web; Programming language","score_opus":0.006260812390844025,"score_gpt":0.23455396375194848,"score_spread":0.22829315136110445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107142491","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.959258,0.004143002,0.0058087567,0.008681091,0.000100893565,0.00025843485,0.000897087,0.00043738482,0.020415314],"genre_scores_gemma":[0.99379855,0.001058242,0.0025124906,0.00035788803,0.000100622274,0.000120391094,0.0005716158,0.0000717564,0.0014084585],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9592372,0.02280899,0.0048154555,0.0016874492,0.009091389,0.002359506],"domain_scores_gemma":[0.51002234,0.38956004,0.05826689,0.009360022,0.025206193,0.0075845793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02625346,0.00051724556,0.00073047756,0.009207381,0.0016481393,0.004611688,0.0017137448,0.0024333296,0.0065788087],"category_scores_gemma":[0.29920965,0.0008730312,0.0006414714,0.0052060303,0.0016323607,0.008496857,0.004729781,0.0015351094,0.0009525874],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015432446,0.00052571035,0.5011413,0.0032987685,0.0001715642,0.0027408942,0.23366696,0.00094026694,0.002988739,0.0080958,0.011289913,0.23359677],"study_design_scores_gemma":[0.00015005053,0.0014771866,0.660447,0.0028456373,0.00042131785,0.0053488365,0.22132216,0.0042783734,0.003042528,0.012128122,0.088080145,0.00045856583],"about_ca_topic_score_codex":0.004962264,"about_ca_topic_score_gemma":0.0022037947,"teacher_disagreement_score":0.02625346,"about_ca_system_score_codex":0.002512527,"about_ca_system_score_gemma":0.0016280923,"threshold_uncertainty_score":0.13884324},"labels":[],"label_agreement":null},{"id":"W2107321288","doi":"10.1109/icsm.2005.52","title":"Improved tool support for the investigation of duplication in software","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software system; Software; Software engineering; Code (set theory); Software maintenance; Program comprehension; Gene duplication; Set (abstract data type); Software construction; Cloning (programming); Software evolution; Programming language","score_opus":0.020479318459759893,"score_gpt":0.26929387428843743,"score_spread":0.24881455582867754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107321288","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08979205,0.00041302168,0.8490395,0.00097188255,0.000098311415,0.00063359685,0.00055199803,0.05547965,0.0030200244],"genre_scores_gemma":[0.16508435,0.00017230361,0.829683,0.00018532485,0.000065276756,0.00042534727,0.0011145088,0.0020177816,0.001252165],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96345,0.0137332,0.0054047164,0.0031925726,0.013157777,0.0010617271],"domain_scores_gemma":[0.7228567,0.18012674,0.013495118,0.05740023,0.024302354,0.0018189198],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024832973,0.0017762403,0.0024249488,0.007620557,0.00072232995,0.0043419963,0.0048594694,0.0037308217,0.004679599],"category_scores_gemma":[0.15724921,0.0014617549,0.0015255779,0.0045134635,0.0013100716,0.011558406,0.0053046495,0.002820756,0.0019050944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013363172,0.0013473826,0.025759466,0.0017567609,0.00025335155,0.0039296346,0.006258553,0.016077971,0.13479295,0.014242275,0.012454806,0.7817905],"study_design_scores_gemma":[0.0012655327,0.0030228416,0.03506932,0.0020049107,0.00084170944,0.011295145,0.0017550396,0.57750714,0.23151405,0.038153898,0.096610345,0.00096006854],"about_ca_topic_score_codex":0.00058402447,"about_ca_topic_score_gemma":0.0008355726,"teacher_disagreement_score":0.024832973,"about_ca_system_score_codex":0.00061342306,"about_ca_system_score_gemma":0.0022648878,"threshold_uncertainty_score":0.1313309},"labels":[],"label_agreement":null},{"id":"W2107796781","doi":"10.1109/icpc.2011.44","title":"Scalable Automatic Concept Mining from Execution Traces","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; TRACE (psycholinguistics); Program comprehension; Identification (biology); Scalability; Artifact (error); Flexibility (engineering); Process (computing); Latent Dirichlet allocation; Code (set theory); Software maintenance; Artificial intelligence; Precision and recall; Machine learning; Software; Programming language; Software system; Topic model","score_opus":0.034939577843170394,"score_gpt":0.25115205328549256,"score_spread":0.21621247544232217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107796781","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05269597,0.00068168057,0.9292442,0.00034847917,0.00006175653,0.0005080591,0.0029227613,0.012363499,0.0011734781],"genre_scores_gemma":[0.192114,0.00042161957,0.793626,0.00007823194,0.000055523768,0.00074576295,0.01095584,0.00047475583,0.0015283299],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975151,0.00044085263,0.00020283001,0.00077880966,0.00085867866,0.00020376976],"domain_scores_gemma":[0.99200433,0.0048415586,0.0005551936,0.00087556796,0.0014668626,0.00025642442],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018899882,0.0017115462,0.0016173612,0.0065805386,0.0012012103,0.0018799485,0.0032872108,0.001116413,0.0016601573],"category_scores_gemma":[0.012012482,0.00065533305,0.0015460933,0.00503112,0.00058891746,0.0036034815,0.0019002615,0.0016587741,0.0012796298],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004639186,0.0005799135,0.0086454665,0.0007248321,0.00016044866,0.0005004056,0.0007766408,0.048706315,0.021730317,0.0070574423,0.012839198,0.8978152],"study_design_scores_gemma":[0.00005219523,0.00007731529,0.0022686755,0.000050653183,0.000049981667,0.00027671276,0.0003575186,0.9553468,0.01249736,0.022794249,0.006188926,0.0000395144],"about_ca_topic_score_codex":0.0060698423,"about_ca_topic_score_gemma":0.010975229,"teacher_disagreement_score":0.0065805386,"about_ca_system_score_codex":0.0010375659,"about_ca_system_score_gemma":0.003067583,"threshold_uncertainty_score":0.0120690465},"labels":[],"label_agreement":null},{"id":"W2107861046","doi":"10.1109/icpc.2006.20","title":"Dynamic Data Structure Analysis for Java Programs","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Heap (data structure); Garbage collection; Java; Data structure; TRACE (psycholinguistics); Static analysis; Dynamic program analysis; Program analysis; Programming language; External Data Representation; Tracing; Dynamic data; Garbage; Operating system","score_opus":0.023285918554327423,"score_gpt":0.2955118892647845,"score_spread":0.27222597071045707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107861046","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060474716,0.000604493,0.9238263,0.0003111031,0.000035700563,0.000069532965,0.0005485625,0.01019462,0.003934952],"genre_scores_gemma":[0.5134161,0.0007845749,0.47833183,0.00014286018,0.000051364175,0.00021896527,0.0015615997,0.0024720905,0.0030206258],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904543,0.00013417557,0.000077205965,0.00016736348,0.00047535362,0.00010058299],"domain_scores_gemma":[0.9972831,0.0010422387,0.00039922906,0.00057313207,0.0006277365,0.00007456592],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00078711484,0.00045147244,0.00040482017,0.002226605,0.0011090254,0.0017893951,0.0007727539,0.0004224811,0.0013280001],"category_scores_gemma":[0.0049708225,0.00049823255,0.00070646347,0.001448336,0.0010078897,0.002196469,0.0014294112,0.0010323721,0.00037012523],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003580235,0.00019277868,0.02132181,0.00061246305,0.00010348767,0.0006427023,0.0018834629,0.0774941,0.102375425,0.18309404,0.009297952,0.60262376],"study_design_scores_gemma":[0.000064164306,0.00011998129,0.0065182922,0.00022611934,0.000095469506,0.00078227854,0.0003461751,0.6254909,0.121937215,0.19342859,0.050842576,0.00014830202],"about_ca_topic_score_codex":0.0038318357,"about_ca_topic_score_gemma":0.0036697255,"teacher_disagreement_score":0.0038318357,"about_ca_system_score_codex":0.0012605896,"about_ca_system_score_gemma":0.0015284736,"threshold_uncertainty_score":0.009146273},"labels":[],"label_agreement":null},{"id":"W2108004063","doi":"10.1109/wpc.2002.1021344","title":"Relocating XML elements from preprocessed to unprocessed code","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Preprocessor; XML; Programming language; Program comprehension; Source code; Element (criminal law); Character (mathematics); Code (set theory); File format; Information retrieval; Operating system; Set (abstract data type); Software","score_opus":0.022450765706301314,"score_gpt":0.2848369643651776,"score_spread":0.26238619865887625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108004063","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13725999,0.0005855512,0.8172387,0.0004700838,0.0004517754,0.0005487632,0.002372272,0.035735875,0.0053370623],"genre_scores_gemma":[0.17385066,0.00047885306,0.8008287,0.00038821457,0.00009813248,0.0002629645,0.004788537,0.01032296,0.008980982],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99897003,0.00023151046,0.00019096071,0.00028573582,0.00025690228,0.00006486598],"domain_scores_gemma":[0.98482597,0.0059274193,0.0010302394,0.0058872206,0.0021832974,0.00014583208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018873556,0.00074412255,0.00046409422,0.0015569407,0.0006417114,0.0021219137,0.001160477,0.00054433435,0.0062175114],"category_scores_gemma":[0.013255445,0.0007219913,0.00057987194,0.0014783294,0.0009331564,0.0021172403,0.0011085612,0.0015369297,0.0038358727],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016492298,0.00026990252,0.011785609,0.001117835,0.00011095289,0.0024976893,0.004529594,0.009153732,0.24109925,0.02507373,0.01405753,0.68865496],"study_design_scores_gemma":[0.0000588062,0.00043325487,0.005801479,0.00024141879,0.00009409956,0.0013517109,0.0006826503,0.026610574,0.830677,0.011297952,0.12262289,0.00012820173],"about_ca_topic_score_codex":0.0011787837,"about_ca_topic_score_gemma":0.0016779333,"teacher_disagreement_score":0.0062175114,"about_ca_system_score_codex":0.00058200984,"about_ca_system_score_gemma":0.0010190367,"threshold_uncertainty_score":0.020799637},"labels":[],"label_agreement":null},{"id":"W2108379022","doi":"10.1109/icpc.2007.27","title":"Metrics for Measuring the Effectiveness of Decompilers and Obfuscators","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Java; Metric (unit); Source code; Set (abstract data type); Class (philosophy); Code (set theory); Software metric; Programming language; Theoretical computer science; Data mining; Software engineering; Software; Software quality; Software development; Artificial intelligence","score_opus":0.02615127465327688,"score_gpt":0.287205072375141,"score_spread":0.2610537977218641,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108379022","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75473094,0.003240778,0.22046207,0.000326505,0.00013588095,0.0014239666,0.0057660365,0.0038047042,0.010109094],"genre_scores_gemma":[0.8088751,0.00082523574,0.18104361,0.00006712449,0.000038073813,0.0011857402,0.006161559,0.00049391686,0.0013096684],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98264474,0.003343484,0.0027253195,0.000916545,0.009789159,0.0005807835],"domain_scores_gemma":[0.87095064,0.080151446,0.0149092395,0.011625067,0.020456173,0.0019074108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008601666,0.0016779556,0.0012432502,0.012826702,0.0007862988,0.0015419363,0.0013811009,0.0014088892,0.0013751574],"category_scores_gemma":[0.061797425,0.0004138466,0.0009502174,0.007988701,0.0013045985,0.0031790289,0.0013047595,0.00096919964,0.00045887506],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016110579,0.0017273187,0.15282735,0.003317453,0.0014194492,0.0004595679,0.0013770611,0.20478939,0.1520896,0.013193815,0.00891723,0.45827073],"study_design_scores_gemma":[0.00027035773,0.0062712277,0.1728707,0.00037099497,0.0007779255,0.0017055615,0.0011356603,0.44604844,0.34145626,0.012334379,0.016274964,0.00048366282],"about_ca_topic_score_codex":0.0015026717,"about_ca_topic_score_gemma":0.0019582177,"teacher_disagreement_score":0.012826702,"about_ca_system_score_codex":0.0011934919,"about_ca_system_score_gemma":0.0009357263,"threshold_uncertainty_score":0.045490503},"labels":[],"label_agreement":null},{"id":"W2108528511","doi":"","title":"A framework for an adaptive refactoring tool","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Armed Forces","funders":"","keywords":"Code refactoring; Computer science; Process (computing); Code (set theory); Software engineering; Source code; Programming language; Software maintenance; Software; Software development","score_opus":0.04958923476828218,"score_gpt":0.3140671583263945,"score_spread":0.2644779235581123,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108528511","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008382198,0.000096163065,0.9884975,0.00013286812,0.000027878497,0.00013920809,0.00005248288,0.009374019,0.0008417798],"genre_scores_gemma":[0.012939929,0.00013311248,0.9838727,0.0000861,0.00003097877,0.0002382574,0.00035515102,0.0009453814,0.0013983895],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9945755,0.001029322,0.00073293573,0.0010593861,0.0022060778,0.0003967401],"domain_scores_gemma":[0.9927978,0.0025816583,0.00050787505,0.001975801,0.0015903246,0.00054659415],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008354439,0.0017988074,0.0015436052,0.0038267842,0.0014924677,0.0044204723,0.0068219253,0.0038334592,0.004784845],"category_scores_gemma":[0.014157903,0.0019070957,0.0028368211,0.0016152367,0.0024998698,0.005314267,0.004855936,0.0050151595,0.0033580384],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034577478,0.0006884647,0.0038368253,0.0010630101,0.0003146813,0.0024313417,0.002071221,0.07241098,0.030522285,0.29867104,0.024346342,0.563298],"study_design_scores_gemma":[0.00024361187,0.00033425438,0.001077797,0.0007401811,0.00024255729,0.00233501,0.00019165532,0.50736946,0.020164883,0.14575137,0.32116398,0.00038527444],"about_ca_topic_score_codex":0.0051457854,"about_ca_topic_score_gemma":0.0042675575,"teacher_disagreement_score":0.008354439,"about_ca_system_score_codex":0.0016428791,"about_ca_system_score_gemma":0.0033584996,"threshold_uncertainty_score":0.044183016},"labels":[],"label_agreement":null},{"id":"W2108545456","doi":"10.1145/2597073.2597083","title":"Mining questions asked by web developers","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":154,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Intel Corporation","keywords":"Computer science; World Wide Web; Ranking (information retrieval); JavaScript; HTML5; Categorization; Rank (graph theory); Web development; Web mining; Web standards; Web design; Standardization; Data science; Web page; Information retrieval; Artificial intelligence","score_opus":0.013344405564697907,"score_gpt":0.25545844299516174,"score_spread":0.24211403743046384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108545456","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9888749,0.00020977805,0.007794126,0.0003690524,0.00001599722,0.00018364507,0.0009864795,0.00013583769,0.0014301536],"genre_scores_gemma":[0.9761478,0.00018619561,0.017550953,0.00027147323,0.000039828792,0.00043503864,0.0037173056,0.000066372995,0.0015850568],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9865296,0.007292017,0.001064396,0.0018161628,0.0026218097,0.0006760034],"domain_scores_gemma":[0.8156955,0.15181027,0.012115648,0.0038562098,0.01461263,0.001909739],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0092488425,0.0006835509,0.00051435764,0.006599112,0.000961228,0.001874288,0.0009619174,0.0016148123,0.00090405165],"category_scores_gemma":[0.08547284,0.00040633365,0.0004871733,0.0029927606,0.0007354909,0.0025262656,0.0020930166,0.0011308813,0.0005070623],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005299614,0.00048206985,0.763656,0.00079489214,0.00012945208,0.00084514695,0.065364316,0.0020104556,0.015851699,0.0012434875,0.00390789,0.14518462],"study_design_scores_gemma":[0.00008168295,0.0006210151,0.7975982,0.000530592,0.00018235153,0.0013280389,0.08135006,0.051859308,0.023061076,0.005880679,0.037307538,0.00019942869],"about_ca_topic_score_codex":0.0027091391,"about_ca_topic_score_gemma":0.0033278468,"teacher_disagreement_score":0.0092488425,"about_ca_system_score_codex":0.0013645763,"about_ca_system_score_gemma":0.0012003027,"threshold_uncertainty_score":0.04891312},"labels":[],"label_agreement":null},{"id":"W2108653655","doi":"10.1109/fie.2006.322402","title":"A Tool for Automated GUI Program Grading","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Java; Graphical user interface; Graphical user interface testing; Software engineering; Grading (engineering); Programming language; Consistency (knowledge bases); Flexibility (engineering); User interface; Operating system; User interface design; Artificial intelligence; Engineering","score_opus":0.015928015375589136,"score_gpt":0.30030399427229404,"score_spread":0.2843759788967049,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108653655","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00255365,0.00012985452,0.6442345,0.0001360097,0.00012870129,0.0003508395,0.0018498048,0.34696382,0.0036528502],"genre_scores_gemma":[0.049634624,0.00024174289,0.88344175,0.00031932528,0.000114723276,0.0011540891,0.012808289,0.040818404,0.011467084],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9918133,0.0018067118,0.0013175823,0.0013606182,0.0033099195,0.00039190231],"domain_scores_gemma":[0.9778676,0.0095727295,0.001439066,0.0054804496,0.0049806214,0.0006594694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069210203,0.0027847956,0.0017808912,0.006428197,0.000966525,0.0037258146,0.0039335815,0.0018345395,0.028955614],"category_scores_gemma":[0.03672753,0.0019659924,0.0015867935,0.0025468983,0.00064735016,0.0051575038,0.003674502,0.0030863301,0.022919582],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057856255,0.00046231181,0.00294912,0.0006133667,0.00011012519,0.00046187418,0.0004963227,0.004389636,0.01295771,0.008758762,0.12516968,0.84305257],"study_design_scores_gemma":[0.0011819783,0.0008082495,0.008618664,0.00096424134,0.00023175073,0.0039004057,0.00031531055,0.28172556,0.12538402,0.050799873,0.52531874,0.0007512713],"about_ca_topic_score_codex":0.0014495952,"about_ca_topic_score_gemma":0.001499268,"teacher_disagreement_score":0.028955614,"about_ca_system_score_codex":0.00096514507,"about_ca_system_score_gemma":0.0015203456,"threshold_uncertainty_score":0.09686619},"labels":[],"label_agreement":null},{"id":"W2108754848","doi":"10.1142/s0218194006002938","title":"A FRAMEWORK FOR THE PRAGMATIC QUALITY OF Z SPECIFICATIONS","year":2006,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Quality (philosophy); Computer science; Compromise; Software quality; Software engineering; Specification language; Risk analysis (engineering); Formal specification; Software quality control; Software; Programming language; Software development; Business","score_opus":0.03496447777121112,"score_gpt":0.3128131571523523,"score_spread":0.2778486793811412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108754848","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004421379,0.001278714,0.97478515,0.0042358013,0.00013938165,0.0002224103,0.00007415957,0.00026890248,0.014574071],"genre_scores_gemma":[0.19252856,0.0014930253,0.800416,0.0008979038,0.00026262924,0.00088493456,0.00017522766,0.00016744611,0.0031742933],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96546936,0.019249767,0.0034709787,0.0021386342,0.008287302,0.0013839768],"domain_scores_gemma":[0.95293134,0.027619364,0.0042683063,0.006291221,0.007565728,0.001323946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030298876,0.0017846148,0.0012147358,0.007527569,0.004116383,0.012224354,0.0030660974,0.005427837,0.004481908],"category_scores_gemma":[0.05394284,0.0015034403,0.00258785,0.0037393903,0.023851631,0.018114803,0.0073917015,0.0070544803,0.00075857894],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007480083,0.0000064428623,0.00007170782,0.00006515008,0.000005779554,0.000049643277,0.00062705093,0.0012659248,0.00025512275,0.9930661,0.00031433834,0.0042653186],"study_design_scores_gemma":[0.00003278864,0.000049809154,0.00014294105,0.0001877456,0.000020160349,0.00015252098,0.00067518867,0.011470389,0.00054284296,0.9654607,0.021220459,0.000044353765],"about_ca_topic_score_codex":0.0063202465,"about_ca_topic_score_gemma":0.0035119755,"teacher_disagreement_score":0.030298876,"about_ca_system_score_codex":0.0056346636,"about_ca_system_score_gemma":0.0069818096,"threshold_uncertainty_score":0.16023767},"labels":[],"label_agreement":null},{"id":"W2108769867","doi":"10.1145/2597073.2597076","title":"The impact of code review coverage and code review participation on software quality: a case study of the qt, VTK, and ITK projects","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":351,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada","keywords":"Code review; Software quality; Computer science; Software engineering; Team software process; Software; Code (set theory); Software inspection; Quality (philosophy); Software peer review; Process (computing); Software construction; Software quality assurance; Software development; Programming language","score_opus":0.08101323997282231,"score_gpt":0.40729258664642815,"score_spread":0.3262793466736058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108769867","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9978935,0.00022111779,0.00046195483,0.00038333592,0.000004410661,0.00004112798,0.000034697867,0.00001161507,0.0009482133],"genre_scores_gemma":[0.9990183,0.00010706386,0.0005320236,0.000060434402,0.0000117023765,0.0000425064,0.000042133357,0.000011787964,0.00017409648],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.938818,0.0359023,0.0034264452,0.004281252,0.014479607,0.0030925553],"domain_scores_gemma":[0.32933646,0.54262376,0.07877497,0.006854142,0.029987982,0.012422716],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.039240032,0.00050227373,0.0006826938,0.0052532223,0.002783804,0.0038512782,0.0016703536,0.0018825632,0.0013100897],"category_scores_gemma":[0.22812086,0.0005463348,0.00074659934,0.004427373,0.004082089,0.004252317,0.00390028,0.0027830906,0.0002514549],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006119413,0.0015294747,0.8681766,0.0005001659,0.000224004,0.0035413841,0.067938894,0.0019790079,0.0018522756,0.0014414832,0.0012670791,0.050937757],"study_design_scores_gemma":[0.000075778145,0.0014434316,0.94004846,0.00030737752,0.00016393575,0.0019577215,0.043729298,0.0073436834,0.0012042536,0.0008971641,0.0027061554,0.00012279332],"about_ca_topic_score_codex":0.013585173,"about_ca_topic_score_gemma":0.018275704,"teacher_disagreement_score":0.96076,"about_ca_system_score_codex":0.0047082114,"about_ca_system_score_gemma":0.0047328365,"threshold_uncertainty_score":0.20752358},"labels":[],"label_agreement":null},{"id":"W2109063509","doi":"10.1109/wcre.2006.12","title":"An Industrial Case Study of Program Artifacts Viewed During Maintenance Tasks","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artifact (error); Workflow; Task (project management); Key (lock); Process (computing); Software maintenance; Software; Exploratory research; Software engineering; Software development; Human–computer interaction; Artificial intelligence; Programming language; Database; Systems engineering; Engineering","score_opus":0.04223354369730759,"score_gpt":0.311929686922468,"score_spread":0.2696961432251604,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109063509","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9826528,0.00012769077,0.014260199,0.0002982083,0.0000103141365,0.00019018014,0.00027986921,0.00020511415,0.0019757198],"genre_scores_gemma":[0.96820015,0.00012713682,0.029681833,0.00006969528,0.0000140994225,0.00013005761,0.00039888272,0.000060544164,0.0013175632],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9954491,0.0025495682,0.00026148293,0.00053936307,0.0009167362,0.0002837063],"domain_scores_gemma":[0.947532,0.039202105,0.0035656109,0.0054567032,0.0030197953,0.0012237364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005012353,0.0004971604,0.0003782167,0.0019652562,0.002138305,0.0011337877,0.001981765,0.0017261164,0.0012481094],"category_scores_gemma":[0.024085395,0.0005212183,0.00046042394,0.0024342386,0.0010899724,0.0014324555,0.0010125412,0.0013641216,0.00037103033],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002333661,0.012888812,0.38885123,0.0016146962,0.00030487138,0.043924507,0.12293489,0.045315154,0.04624447,0.013882382,0.0121046975,0.30960062],"study_design_scores_gemma":[0.0013068789,0.009373323,0.49003693,0.00055758754,0.00054100144,0.025592048,0.068687126,0.22705741,0.08157028,0.013543471,0.08128219,0.00045175833],"about_ca_topic_score_codex":0.009724287,"about_ca_topic_score_gemma":0.019215886,"teacher_disagreement_score":0.009724287,"about_ca_system_score_codex":0.0014587917,"about_ca_system_score_gemma":0.0013585603,"threshold_uncertainty_score":0.026508152},"labels":[],"label_agreement":null},{"id":"W2109125971","doi":"10.1145/585058.585065","title":"The relevance of software documentation, tools and technologies","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":261,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"LG Display","keywords":"Documentation; Software documentation; Relevance (law); Computer science; Software engineering; Internal documentation; Technical documentation; Software; Software maintenance; Automation; Software development; Knowledge management; Data science; World Wide Web; Software construction; Engineering; Programming language","score_opus":0.021083646717175407,"score_gpt":0.25959428714179356,"score_spread":0.23851064042461814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109125971","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9808456,0.0023369663,0.0025749004,0.0022959404,0.00003581867,0.000055720095,0.00005506768,0.000031430827,0.01176854],"genre_scores_gemma":[0.99634343,0.00076829595,0.001568668,0.00039320736,0.000037671718,0.00002239788,0.000059621372,0.000015282827,0.0007915598],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.96859276,0.014967444,0.0029736285,0.0009769378,0.011331924,0.0011574501],"domain_scores_gemma":[0.83256876,0.11908057,0.021545807,0.0038211525,0.018003095,0.0049805893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017651003,0.0002166308,0.0003890692,0.0034065358,0.0017528214,0.0038649696,0.0004020536,0.0012300957,0.0018899019],"category_scores_gemma":[0.14672597,0.00031975732,0.0003298576,0.0025003501,0.0015602123,0.004125872,0.002356615,0.0013119702,0.00035558402],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004602269,0.00029198616,0.61217767,0.0017467596,0.00015294642,0.001201517,0.11990836,0.00058910315,0.0068821986,0.00412616,0.0031712533,0.24929191],"study_design_scores_gemma":[0.000051479186,0.0012732719,0.7485338,0.0016245833,0.00018493852,0.004444783,0.18630567,0.0017502623,0.002312018,0.006758057,0.0466049,0.00015619896],"about_ca_topic_score_codex":0.0014977582,"about_ca_topic_score_gemma":0.00238422,"teacher_disagreement_score":0.017651003,"about_ca_system_score_codex":0.0013027982,"about_ca_system_score_gemma":0.0021315739,"threshold_uncertainty_score":0.09334856},"labels":[],"label_agreement":null},{"id":"W2109156518","doi":"10.1007/s10664-013-9258-8","title":"Bug characteristics in open source software","year":2013,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":232,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Science Foundation","keywords":"Software bug; Computer science; Concurrency; Security bug; Software regression; Linux kernel; Operating system; Software; Source code; Software engineering; Software system; Software security assurance; Software construction","score_opus":0.025828708118874232,"score_gpt":0.2819716111627644,"score_spread":0.2561429030438902,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109156518","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99893266,0.0001388409,0.0004897184,0.00004809863,0.0000025925658,0.000005336566,0.0000654107,0.0000150928045,0.000302247],"genre_scores_gemma":[0.99939394,0.000033478784,0.00027067083,0.0000074245154,0.0000045486067,0.000005197378,0.00012062556,0.000014804597,0.00014937719],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99597484,0.0009827387,0.00063633564,0.0005358008,0.0015050089,0.00036519975],"domain_scores_gemma":[0.7644025,0.13662311,0.072840855,0.0068358267,0.014526241,0.0047714706],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0039257198,0.00024745148,0.00032258962,0.006734596,0.00058761856,0.0014316371,0.0006364435,0.0010092532,0.0014879977],"category_scores_gemma":[0.101130396,0.0004430387,0.0005419804,0.00479579,0.0010587752,0.0029916924,0.0013783533,0.001157922,0.0002504782],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012357776,0.00010299069,0.9868183,0.00002856628,0.00003950317,0.0000933248,0.0007917371,0.0005870449,0.0005107227,0.0004694743,0.00018347753,0.0102512585],"study_design_scores_gemma":[0.000008939404,0.00012793357,0.9945679,0.000028121292,0.00003200799,0.00034225095,0.0008201026,0.0025267939,0.00026563898,0.0010049215,0.00026045236,0.000015007193],"about_ca_topic_score_codex":0.0037961216,"about_ca_topic_score_gemma":0.006649665,"teacher_disagreement_score":0.99607426,"about_ca_system_score_codex":0.00093364954,"about_ca_system_score_gemma":0.00084224535,"threshold_uncertainty_score":0.02076149},"labels":[],"label_agreement":null},{"id":"W2109436601","doi":"10.1109/csmr.2012.19","title":"Understanding Structural Complexity Evolution: A Quantitative Analysis","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of British Columbia","keywords":"Commit; Structural complexity; Maintainability; Computer science; Software evolution; Variation (astronomy); Programming complexity; Software; Complexity management; Source code; Software development; Data science; Software engineering; Database; Software construction; Artificial intelligence; Business; Programming language","score_opus":0.2623030571746268,"score_gpt":0.3594013869041854,"score_spread":0.09709832972955856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109436601","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.83996475,0.00046941772,0.14774431,0.0006459475,0.000018427816,0.0005504988,0.0008796337,0.00021665232,0.009510447],"genre_scores_gemma":[0.97622275,0.000083737665,0.022744844,0.000026133805,0.000011307475,0.00028997325,0.0002862893,0.000036960937,0.0002979914],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9888053,0.00470947,0.00063159293,0.0010326059,0.004472945,0.0003480833],"domain_scores_gemma":[0.80040973,0.15712288,0.01855407,0.0054641264,0.016785232,0.0016639347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011297962,0.00062986335,0.00043771625,0.007820094,0.00071521057,0.002481251,0.0008301586,0.00066074217,0.0031178899],"category_scores_gemma":[0.06914199,0.0003559829,0.00087520416,0.004220565,0.0026917616,0.004246492,0.0016536542,0.0009718529,0.00024753713],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005330282,0.00071180484,0.7205682,0.0021929971,0.00073472987,0.0003797179,0.028454514,0.023291994,0.023201993,0.030067045,0.0018870134,0.16797702],"study_design_scores_gemma":[0.00006732557,0.0012372095,0.78947943,0.00037510635,0.00037319024,0.000769454,0.017193614,0.13126618,0.012791449,0.03807274,0.008178466,0.00019585264],"about_ca_topic_score_codex":0.0014384764,"about_ca_topic_score_gemma":0.0011145837,"teacher_disagreement_score":0.011297962,"about_ca_system_score_codex":0.0021488718,"about_ca_system_score_gemma":0.0015872999,"threshold_uncertainty_score":0.05975002},"labels":[],"label_agreement":null},{"id":"W2109548412","doi":"10.1109/wcre.2012.55","title":"An Empirical Study of the Effect of File Editing Patterns on Software Quality","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software; Quartile; Software bug; Quality (philosophy); Software quality; Empirical research; World Wide Web; Software development; Database; Operating system","score_opus":0.03710189101573183,"score_gpt":0.36967992783199405,"score_spread":0.3325780368162622,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109548412","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99809974,0.00011264663,0.0005971967,0.0001151852,0.000003358536,0.000023793691,0.00021863982,0.000028064604,0.0008013648],"genre_scores_gemma":[0.99898404,0.00004947833,0.0005321058,0.000019398272,0.0000072629496,0.000019264411,0.000236599,0.000011586287,0.0001403025],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98246276,0.0067914,0.0018015078,0.0023405852,0.0053048935,0.0012988353],"domain_scores_gemma":[0.40289903,0.45197657,0.101867884,0.015792299,0.021028865,0.0064353812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011083648,0.00043808942,0.00040897654,0.0032574371,0.00066251337,0.0019507674,0.0013030575,0.0009952191,0.0022077847],"category_scores_gemma":[0.17477861,0.00039839742,0.0007548634,0.0051291403,0.0017373536,0.0043330137,0.0014395557,0.0019678692,0.00046373293],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015452711,0.00023828179,0.98908305,0.000061087434,0.00012337341,0.00011917493,0.00087040337,0.00074020517,0.00040174773,0.00009824233,0.00022315806,0.007886758],"study_design_scores_gemma":[0.000008256607,0.00021175056,0.99642026,0.000011947585,0.00002434074,0.00013660527,0.00091682666,0.0017112544,0.00025073325,0.0000951099,0.00020129142,0.000011617395],"about_ca_topic_score_codex":0.0044704424,"about_ca_topic_score_gemma":0.0043340838,"teacher_disagreement_score":0.011083648,"about_ca_system_score_codex":0.001103062,"about_ca_system_score_gemma":0.0007597781,"threshold_uncertainty_score":0.05861664},"labels":[],"label_agreement":null},{"id":"W2109580177","doi":"10.1016/j.scico.2013.11.027","title":"An insight into the dispersion of changes in cloned and non-cloned code: A genealogy based empirical study","year":2013,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Code (set theory); Cloning (programming); Empirical research; Genealogy; Programming language; History; Statistics; Mathematics","score_opus":0.02139409513525516,"score_gpt":0.31069825172466453,"score_spread":0.2893041565894094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109580177","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99566615,0.00018320017,0.0018258265,0.00018455938,0.0000024397225,0.0000060473844,0.00010584567,0.000014037491,0.002011688],"genre_scores_gemma":[0.9989592,0.00005988957,0.00045223685,0.00002020796,0.0000041655417,0.000003891011,0.00008835486,0.000010566934,0.00040138402],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9973538,0.0011924125,0.00016871539,0.0005744431,0.0005223053,0.00018834058],"domain_scores_gemma":[0.9056657,0.073942915,0.010609125,0.0058490057,0.0026338894,0.0012993434],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004982007,0.000121295394,0.00039330215,0.004113319,0.0010331159,0.0017796744,0.0011322211,0.001047466,0.004116815],"category_scores_gemma":[0.07162661,0.00029961806,0.0002501338,0.0043991497,0.0035458528,0.004208814,0.0015793398,0.001182401,0.0004101551],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024255224,0.00017079667,0.9300351,0.000054180106,0.00015783716,0.0004749513,0.012144018,0.0050425977,0.0019301532,0.027079543,0.0005389222,0.022129245],"study_design_scores_gemma":[0.000027739168,0.00028067234,0.92292476,0.00006397406,0.00008734792,0.0021077686,0.0124132745,0.023481254,0.0013374655,0.0339842,0.0032283121,0.00006328129],"about_ca_topic_score_codex":0.0060214894,"about_ca_topic_score_gemma":0.005679198,"teacher_disagreement_score":0.995018,"about_ca_system_score_codex":0.00091278634,"about_ca_system_score_gemma":0.0003869108,"threshold_uncertainty_score":0.026347697},"labels":[],"label_agreement":null},{"id":"W2109773497","doi":"10.1109/wcre.2000.891464","title":"Towards portable source code representations using XML","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Programming language; XML; Source code; Static program analysis; Code generation; KPI-driven code analysis; Markup language; Domain-specific language; Software portability; Abstract syntax; Software engineering; Domain (mathematical analysis); Software; Software development; Semantics (computer science); Operating system","score_opus":0.0672649429835274,"score_gpt":0.31187057572451615,"score_spread":0.24460563274098873,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109773497","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002558031,0.00009024062,0.99000937,0.00022644339,0.000032785974,0.000070947564,0.00021364588,0.0051823645,0.0016162335],"genre_scores_gemma":[0.030626582,0.00056892313,0.96147096,0.00013922955,0.00003823518,0.00023628699,0.001708855,0.0017702944,0.0034407452],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976036,0.00077749137,0.00033764538,0.00025276066,0.0009194908,0.00010902974],"domain_scores_gemma":[0.99275994,0.00215026,0.0006885873,0.002661387,0.0016042436,0.0001356159],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003920279,0.000813391,0.0005550232,0.0032863924,0.000664246,0.005139609,0.002060042,0.0016893672,0.0029381798],"category_scores_gemma":[0.013285028,0.00092888117,0.0010558782,0.0032351527,0.0011477941,0.0054179374,0.0028847228,0.0025995849,0.002451929],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019172332,0.0001794924,0.0015966127,0.00067125674,0.00008099308,0.0012234694,0.0030825615,0.023708373,0.028468322,0.47374594,0.01749801,0.44955316],"study_design_scores_gemma":[0.00013009907,0.00012416071,0.0010313594,0.0008092126,0.000107103566,0.0013084507,0.00088461034,0.28452554,0.07552181,0.30578682,0.32964265,0.00012809147],"about_ca_topic_score_codex":0.0010143525,"about_ca_topic_score_gemma":0.0008716584,"teacher_disagreement_score":0.005139609,"about_ca_system_score_codex":0.00065424194,"about_ca_system_score_gemma":0.0011568773,"threshold_uncertainty_score":0.020732641},"labels":[],"label_agreement":null},{"id":"W2109792733","doi":"10.1109/wpc.2002.1021305","title":"Mining system-user interaction traces for use case models","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Program comprehension; User interface; User modeling; Set (abstract data type); Human–computer interaction; Source code; Code (set theory); Business process reengineering; Matching (statistics); Legacy system; Data mining; Programming language; Software system; Software","score_opus":0.07050222923159072,"score_gpt":0.3051729161856184,"score_spread":0.2346706869540277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109792733","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23998727,0.00048499924,0.74083734,0.00068514375,0.00003873834,0.0014203484,0.007748618,0.0056256014,0.0031719077],"genre_scores_gemma":[0.468784,0.00041581714,0.5107364,0.000062225845,0.00002263279,0.0014418858,0.016808733,0.00037412747,0.0013541009],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99506605,0.0015379594,0.0006240408,0.00074474374,0.0018107296,0.00021644609],"domain_scores_gemma":[0.9706575,0.019018281,0.0022012154,0.004951707,0.0028121732,0.00035909907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044039874,0.0016024321,0.0008705943,0.006954552,0.00079605,0.0028583207,0.0024118677,0.001740287,0.0017305928],"category_scores_gemma":[0.043315724,0.0010180963,0.0019341321,0.0035404155,0.00079628057,0.0033339302,0.0014833933,0.0015293785,0.00086022704],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004967888,0.0017084236,0.10580061,0.0015994713,0.0005453908,0.003244247,0.00474964,0.41716614,0.016929988,0.028465208,0.009508007,0.4097861],"study_design_scores_gemma":[0.000029288207,0.00010391268,0.0064229392,0.00010926335,0.00007959674,0.00040657123,0.0004573931,0.9679828,0.0067264084,0.0117958,0.0058334945,0.00005253428],"about_ca_topic_score_codex":0.011809954,"about_ca_topic_score_gemma":0.01606403,"teacher_disagreement_score":0.011809954,"about_ca_system_score_codex":0.0020044395,"about_ca_system_score_gemma":0.0031265349,"threshold_uncertainty_score":0.023482382},"labels":[],"label_agreement":null},{"id":"W2109942136","doi":"10.1007/s10664-008-9104-6","title":"A study of the non-linear adjustment for analogy based software cost estimation","year":2009,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Saskatchewan; University of Wisconsin-Madison","keywords":"Categorical variable; Analogy; Flexibility (engineering); Computer science; Software; Estimation; Artificial intelligence; Artificial neural network; Machine learning; Data mining; Statistics; Mathematics; Engineering","score_opus":0.035501370270765505,"score_gpt":0.320520460559485,"score_spread":0.2850190902887195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109942136","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035159968,0.00051997515,0.95855,0.00050037197,0.000088651024,0.000075555174,0.000031784883,0.00024266173,0.004831104],"genre_scores_gemma":[0.68609124,0.0003858065,0.30589908,0.0001916173,0.00011551115,0.00011301863,0.0001029605,0.00024172482,0.0068590697],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99540544,0.0029443156,0.00014279233,0.00052336045,0.0008295633,0.00015455589],"domain_scores_gemma":[0.9700991,0.02481494,0.0012163578,0.0023495746,0.001368428,0.00015154356],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005097798,0.00050176005,0.0006935793,0.000860425,0.00052389153,0.0014218342,0.0027969424,0.00094254775,0.0061103893],"category_scores_gemma":[0.07071627,0.00048785386,0.00087550655,0.0022467843,0.0009844913,0.003252907,0.0012339397,0.0024260634,0.00049672223],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024447666,0.00025242683,0.0061567025,0.00042054808,0.00020593038,0.00022558746,0.0005378405,0.26642486,0.004636369,0.38619485,0.002366531,0.33233383],"study_design_scores_gemma":[0.000020011763,0.00012290193,0.0033019164,0.000027221266,0.000050299397,0.00013524215,0.000078174206,0.89791447,0.0016461434,0.09288818,0.0037771077,0.000038271457],"about_ca_topic_score_codex":0.0037545753,"about_ca_topic_score_gemma":0.002800257,"teacher_disagreement_score":0.0061103893,"about_ca_system_score_codex":0.0010721005,"about_ca_system_score_gemma":0.0010012792,"threshold_uncertainty_score":0.026960015},"labels":[],"label_agreement":null},{"id":"W2109961432","doi":"10.1109/esem.2013.47","title":"Towards a Metric Suite Proposal to Quantify Confirmation Biases of Developers","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software metric; Computer science; Metric (unit); Suite; Software quality; Software; Process (computing); Quality (philosophy); Software development; Software measurement; Product metric; Empirical research; Set (abstract data type); Data science; Data mining; Software engineering; Engineering; Metric space; Mathematics; Statistics","score_opus":0.045491905320010385,"score_gpt":0.3075101891520013,"score_spread":0.2620182838319909,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109961432","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028438937,0.00070847286,0.9604411,0.0017087756,0.00017289766,0.0010825796,0.0010047146,0.0012842888,0.0051582474],"genre_scores_gemma":[0.15627718,0.0003153387,0.8386716,0.00020354982,0.0000891031,0.0024070083,0.0013561196,0.00015733129,0.0005227991],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.93435496,0.03481562,0.008719591,0.0039085443,0.01710859,0.0010927987],"domain_scores_gemma":[0.8444766,0.06840261,0.017652933,0.014130739,0.052401405,0.0029358251],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04607869,0.0030507338,0.0022914722,0.01573889,0.0019041635,0.0078950655,0.0038869497,0.0030077999,0.0014738862],"category_scores_gemma":[0.18408418,0.000929737,0.0020894203,0.011765634,0.002849171,0.009840378,0.0052305185,0.003225713,0.0011362636],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046017865,0.0009181696,0.09246801,0.0022438464,0.0009275285,0.00035242783,0.004464852,0.046819028,0.013164961,0.2308,0.013375294,0.59400564],"study_design_scores_gemma":[0.00023988287,0.0036435202,0.05182869,0.0014107629,0.0006711744,0.0018688095,0.003958911,0.42942694,0.02252704,0.4147449,0.068914935,0.00076445896],"about_ca_topic_score_codex":0.0021424375,"about_ca_topic_score_gemma":0.0016881885,"teacher_disagreement_score":0.9539213,"about_ca_system_score_codex":0.0035609624,"about_ca_system_score_gemma":0.005418572,"threshold_uncertainty_score":0.24369037},"labels":[],"label_agreement":null},{"id":"W2110374486","doi":"10.1145/1806799.1806872","title":"Summarizing software artifacts","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":197,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software; Conversation; Software bug; Focus (optics); Generator (circuit theory); Software engineering; Software development; Software maintenance; Programming language; Power (physics)","score_opus":0.014433922352686823,"score_gpt":0.25335965784768155,"score_spread":0.23892573549499474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110374486","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14712438,0.0020656895,0.8232171,0.0013863728,0.0005462816,0.0011816326,0.003593698,0.011558635,0.0093261795],"genre_scores_gemma":[0.41791221,0.0013714816,0.5613977,0.00031223765,0.0003309129,0.0004709415,0.0105307335,0.0010322467,0.0066416],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972699,0.0013483913,0.00024389283,0.00047237205,0.00059024635,0.0000750696],"domain_scores_gemma":[0.98273957,0.009642593,0.0013433397,0.0029955343,0.0029600705,0.0003188214],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0034509778,0.001127719,0.0005878822,0.00266877,0.00077968556,0.0016507022,0.0010534001,0.00083195337,0.003932952],"category_scores_gemma":[0.030398654,0.0003489336,0.0006779029,0.0017314061,0.00034966806,0.0020598068,0.0012414916,0.0007267299,0.0018120972],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005281276,0.0003125077,0.011282156,0.0014436512,0.00018770652,0.00059471134,0.0041084792,0.02046573,0.031556323,0.013916657,0.026175717,0.8894282],"study_design_scores_gemma":[0.0003709349,0.003303869,0.029765975,0.0009624981,0.0014231408,0.003747428,0.0056115924,0.40480086,0.12371516,0.11955413,0.3062988,0.00044565025],"about_ca_topic_score_codex":0.0008942974,"about_ca_topic_score_gemma":0.0014803656,"teacher_disagreement_score":0.996549,"about_ca_system_score_codex":0.0004186829,"about_ca_system_score_gemma":0.0009104051,"threshold_uncertainty_score":0.018250704},"labels":[],"label_agreement":null},{"id":"W2110419873","doi":"10.1007/s00766-003-0177-x","title":"More requirements engineering adventures with building contractors","year":2003,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Viewpoints; Adventure; Section (typography); Requirements engineering; Engineering ethics; Engineering; Special section; Engineering management; Computer science; Artificial intelligence; Visual arts; Art; Engineering physics; Software; Programming language","score_opus":0.017839099182069613,"score_gpt":0.26782713997204066,"score_spread":0.24998804078997106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110419873","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03912703,0.014422485,0.121455275,0.32320452,0.019912763,0.0001614721,0.00031943267,0.0021317906,0.4792652],"genre_scores_gemma":[0.22890306,0.008490969,0.1264632,0.03985629,0.012749734,0.000155385,0.00062243,0.0032814716,0.5794775],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.97867614,0.010184412,0.00066987216,0.001932716,0.0069582486,0.0015785452],"domain_scores_gemma":[0.9518748,0.015177187,0.0021095797,0.011091628,0.009542869,0.0102039315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022251561,0.0010422554,0.00069225434,0.001959438,0.006222177,0.012045179,0.00287769,0.005967721,0.074414805],"category_scores_gemma":[0.039422292,0.0011418457,0.0016821712,0.002920824,0.0043240143,0.021749714,0.0074895155,0.013521765,0.011975019],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029177326,0.0009849776,0.0028777167,0.0004509569,0.00007389312,0.0006622401,0.012635549,0.0017888464,0.0047464287,0.23689859,0.38015747,0.35843152],"study_design_scores_gemma":[0.000022765526,0.00014060784,0.0012294628,0.00019038942,0.00001322172,0.0005347369,0.0036941105,0.00081199,0.0010746425,0.029341698,0.9628934,0.000052902364],"about_ca_topic_score_codex":0.003245681,"about_ca_topic_score_gemma":0.009638084,"teacher_disagreement_score":0.074414805,"about_ca_system_score_codex":0.0045634294,"about_ca_system_score_gemma":0.0033193864,"threshold_uncertainty_score":0.24894232},"labels":[],"label_agreement":null},{"id":"W2110431240","doi":"10.1109/tse.2002.1019484","title":"Assessing the applicability of fault-proneness models across object-oriented software projects","year":2002,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":351,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Universidade Federal do Rio Grande do Sul; National Science Foundation","keywords":"Computer science; Software metric; Software system; Software; Data mining; A priori and a posteriori; Multivariate adaptive regression splines; Software development; Java; Fault (geology); Software fault tolerance; Software engineering; Machine learning; Software quality; Regression analysis; Programming language","score_opus":0.04020972332741948,"score_gpt":0.2891038020238414,"score_spread":0.2488940786964219,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110431240","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9522756,0.00025857575,0.04520004,0.0002831032,0.000011894052,0.00011058561,0.00036464032,0.00019108976,0.0013044978],"genre_scores_gemma":[0.98904115,0.00008385528,0.010141595,0.00002398865,0.000010173352,0.00007600174,0.00043562838,0.000027696244,0.00016005228],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.989421,0.0059182155,0.0006915259,0.0014149982,0.002021598,0.00053267594],"domain_scores_gemma":[0.84448105,0.1212119,0.015718376,0.010952653,0.0062187854,0.0014171948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026272558,0.0011534463,0.00091653765,0.0054267147,0.0005754826,0.0016504949,0.002134443,0.0018882575,0.0008983266],"category_scores_gemma":[0.10249006,0.0005668817,0.0017470713,0.0034022925,0.0011576964,0.003849548,0.002505996,0.0018406687,0.00024460215],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075951003,0.00085127715,0.41136283,0.00017987624,0.00087060314,0.00017501798,0.0013263939,0.51517606,0.0014051794,0.0045370953,0.0005568232,0.06279925],"study_design_scores_gemma":[0.00003400454,0.0014241135,0.15601353,0.00004993236,0.00012369768,0.00011929553,0.000786612,0.8304014,0.00082289823,0.009556324,0.00060949,0.000058628357],"about_ca_topic_score_codex":0.0059622442,"about_ca_topic_score_gemma":0.004746999,"teacher_disagreement_score":0.026272558,"about_ca_system_score_codex":0.0017533789,"about_ca_system_score_gemma":0.00080031465,"threshold_uncertainty_score":0.13894421},"labels":[],"label_agreement":null},{"id":"W2110499788","doi":"10.1109/wcre.2012.63","title":"OpenTrace: An Open Source Workbench for Automatic Software Traceability Link Recovery","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Workbench; Computer science; TRACE (psycholinguistics); Traceability; Artifact (error); Preprocessor; Software; Open source; Open source software; Plug-in; Link (geometry); Software engineering; Data mining; Operating system; Programming language; Artificial intelligence; Visualization","score_opus":0.04144076259156273,"score_gpt":0.32083163698372047,"score_spread":0.27939087439215776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110499788","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021398453,0.0005046756,0.4792739,0.00035587797,0.00035583146,0.0012875366,0.028790332,0.46382335,0.0042101126],"genre_scores_gemma":[0.10474621,0.0010119246,0.60505944,0.0004423238,0.000237498,0.003833464,0.18834378,0.08890024,0.0074251583],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9945825,0.0010448915,0.00094243325,0.0008466338,0.0022935092,0.00029007567],"domain_scores_gemma":[0.9715332,0.014397201,0.0015850473,0.008391944,0.003334987,0.00075756386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067733624,0.0040707504,0.0016156819,0.009583833,0.0011680619,0.0035945093,0.007134992,0.0018415055,0.015004618],"category_scores_gemma":[0.033872727,0.0015848835,0.002043818,0.0055166017,0.0009874847,0.0052873464,0.004764206,0.0030415575,0.007534437],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023605102,0.0032196543,0.011080114,0.006008633,0.0015086791,0.0018584521,0.0019026219,0.049399313,0.05395673,0.013442909,0.2778381,0.57742435],"study_design_scores_gemma":[0.001843826,0.001968563,0.0149291735,0.0011539249,0.0005458178,0.0018739251,0.00097468967,0.46407086,0.17094612,0.02980693,0.31087932,0.0010068666],"about_ca_topic_score_codex":0.004693908,"about_ca_topic_score_gemma":0.0044693663,"teacher_disagreement_score":0.015004618,"about_ca_system_score_codex":0.00080467685,"about_ca_system_score_gemma":0.0026500823,"threshold_uncertainty_score":0.050195396},"labels":[],"label_agreement":null},{"id":"W2110617721","doi":"10.1109/icse.2009.5070565","title":"SemDiff: Analysis and recommendation support for API evolution","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Context (archaeology); Task (project management); Software engineering; Software evolution; Application programming interface; Programming language; Software development; Software; Systems engineering; Engineering; Software construction","score_opus":0.014837028268179139,"score_gpt":0.28701308743134185,"score_spread":0.27217605916316273,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110617721","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0150387585,0.00050380005,0.5825068,0.0007015471,0.00010822895,0.0004484058,0.0087665,0.38872653,0.0031994968],"genre_scores_gemma":[0.09015318,0.00038921388,0.8644651,0.00046435808,0.000095800584,0.00054634357,0.026500426,0.011404883,0.0059807403],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9945312,0.0008842839,0.0008119847,0.001338277,0.0022329283,0.00020136153],"domain_scores_gemma":[0.9719016,0.012852606,0.002360383,0.006894824,0.005146453,0.000844027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075250934,0.0024289098,0.0014817188,0.01000235,0.0012977466,0.0034327363,0.0041111573,0.0020243851,0.0100316815],"category_scores_gemma":[0.038884602,0.0018045015,0.0015457754,0.00425888,0.00052652525,0.007040852,0.0029129232,0.0021887866,0.0057399296],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011500366,0.0004167699,0.023011306,0.00080223073,0.00033108704,0.0007152012,0.0009292945,0.00771695,0.012632373,0.005716008,0.11720006,0.82937866],"study_design_scores_gemma":[0.00045221313,0.00031849803,0.01176494,0.00035735904,0.00025056707,0.0011941928,0.00048614826,0.6980679,0.0591637,0.023530934,0.20400994,0.0004035466],"about_ca_topic_score_codex":0.009849015,"about_ca_topic_score_gemma":0.019177007,"teacher_disagreement_score":0.0100316815,"about_ca_system_score_codex":0.0012182501,"about_ca_system_score_gemma":0.0018855814,"threshold_uncertainty_score":0.03979695},"labels":[],"label_agreement":null},{"id":"W2110661833","doi":"10.1109/cmpsac.2002.1044543","title":"A fuzzy logic framework to improve the performance and interpretation of rule-based quality prediction models for OO software","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Université de Montréal","funders":"","keywords":"Computer science; Antecedent (behavioral psychology); Fuzzy logic; Software quality; Data mining; Artificial intelligence; Software; Machine learning; Software engineering; Software development; Programming language","score_opus":0.025314773109635248,"score_gpt":0.29203431696264764,"score_spread":0.2667195438530124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110661833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0044591026,0.00017202352,0.9930016,0.00023011716,0.000041762185,0.000049468556,0.000060697654,0.00040973717,0.0015756143],"genre_scores_gemma":[0.13934702,0.00025733034,0.8587522,0.00019080647,0.00009161161,0.00015174688,0.00019683648,0.00004489995,0.00096767285],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99897754,0.0002809549,0.00010372832,0.00015560436,0.00042716932,0.00005498577],"domain_scores_gemma":[0.99732023,0.0015959631,0.00017542967,0.00022631168,0.0006308351,0.000051183542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003786848,0.00082867156,0.0007195124,0.0013511505,0.00084359274,0.0020743397,0.0017284859,0.0011936278,0.002258289],"category_scores_gemma":[0.0082316045,0.00039361478,0.0013229097,0.0010472608,0.0009848202,0.0020590287,0.0009560293,0.0019157879,0.0004576609],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019841174,0.00036413528,0.0015850189,0.00027702097,0.00022574759,0.00045208915,0.00061446463,0.43341583,0.008673115,0.2194695,0.0036804113,0.3310442],"study_design_scores_gemma":[0.000028984787,0.000049888164,0.00020512084,0.00005500152,0.00005687424,0.000058575846,0.00002401938,0.93784183,0.0019294503,0.056584515,0.0031434572,0.000022276774],"about_ca_topic_score_codex":0.012744684,"about_ca_topic_score_gemma":0.01039669,"teacher_disagreement_score":0.012744684,"about_ca_system_score_codex":0.0015602324,"about_ca_system_score_gemma":0.0017295122,"threshold_uncertainty_score":0.025341034},"labels":[],"label_agreement":null},{"id":"W2110674601","doi":"10.1109/wpc.2003.1199207","title":"Scaling an object-oriented system execution visualizer through sampling","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Animation; TRACE (psycholinguistics); Visualization; Software visualization; Task (project management); Software; Object-oriented programming; Software system; Human–computer interaction; Software engineering; Component-based software engineering; Programming language; Computer graphics (images); Data mining; Systems engineering","score_opus":0.03639352949789695,"score_gpt":0.3271955682052205,"score_spread":0.2908020387073236,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110674601","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26750848,0.00033805455,0.6922271,0.0006153387,0.00009380705,0.00026664766,0.00038095293,0.032261427,0.00630822],"genre_scores_gemma":[0.55921113,0.0003876777,0.43551812,0.00015897112,0.000036868576,0.00022845778,0.0004845893,0.0021989772,0.0017751678],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99950695,0.00015537681,0.000034558307,0.00007643784,0.0001857304,0.00004089082],"domain_scores_gemma":[0.99593157,0.0025770646,0.00022057637,0.00067609304,0.00037301402,0.00022160828],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016643474,0.0007898359,0.00039595642,0.00086637225,0.00034347284,0.0013123032,0.0009892818,0.00056970096,0.0023372497],"category_scores_gemma":[0.006624741,0.00044695166,0.00047183505,0.000504871,0.00047357797,0.0015973268,0.0017193255,0.00096241565,0.0004218553],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023119347,0.0007554475,0.017364684,0.000625478,0.00015695531,0.00095947285,0.0062934537,0.17257456,0.22385032,0.017224403,0.012408372,0.54547495],"study_design_scores_gemma":[0.00019844224,0.000371361,0.0047052656,0.00007165141,0.000059639508,0.000345153,0.00028938623,0.90858096,0.054778315,0.008064683,0.022445485,0.00008960843],"about_ca_topic_score_codex":0.0017366755,"about_ca_topic_score_gemma":0.002250744,"teacher_disagreement_score":0.0023372497,"about_ca_system_score_codex":0.00051472225,"about_ca_system_score_gemma":0.00043046824,"threshold_uncertainty_score":0.008801997},"labels":[],"label_agreement":null},{"id":"W2110847937","doi":"10.1016/j.scico.2011.11.002","title":"Tuning research tools for scalability and performance: The NiCad experience","year":2011,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Queen's University","funders":"Center for Advanced Study, University of Illinois at Urbana-Champaign; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Scalability; Laptop; Factor (programming language); Software; Subsequence; Parsing; Code (set theory); Process (computing); Software engineering; Programming language; Operating system","score_opus":0.1341347190317785,"score_gpt":0.35838879632034276,"score_spread":0.22425407728856425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110847937","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16450082,0.009401965,0.6364604,0.014513688,0.0012541476,0.0005832776,0.0011716064,0.055772573,0.11634148],"genre_scores_gemma":[0.37042323,0.003126683,0.60117084,0.0014393822,0.0004303636,0.00048457432,0.0013025715,0.007129748,0.014492679],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98460567,0.0071416916,0.0006432152,0.0013479567,0.0056928326,0.0005685963],"domain_scores_gemma":[0.9007604,0.05447537,0.0016964885,0.020881549,0.018912178,0.0032740287],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029497571,0.0012389225,0.0012965945,0.0029383667,0.0015973403,0.0054567354,0.004926632,0.0013938859,0.0074742325],"category_scores_gemma":[0.07485249,0.0009397957,0.00066017435,0.003426273,0.0023722472,0.013147299,0.0041901073,0.0054500992,0.002072908],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027428246,0.0012326195,0.011239152,0.00067423045,0.00017352054,0.00015397719,0.0026707717,0.01656108,0.015134665,0.1374611,0.07371888,0.7382372],"study_design_scores_gemma":[0.0015753108,0.0013192414,0.008617536,0.000801935,0.0004303589,0.0007482729,0.0022806493,0.34874797,0.075280845,0.20429759,0.35545105,0.00044918153],"about_ca_topic_score_codex":0.0041169263,"about_ca_topic_score_gemma":0.0041651563,"teacher_disagreement_score":0.029497571,"about_ca_system_score_codex":0.0026014615,"about_ca_system_score_gemma":0.0041055502,"threshold_uncertainty_score":0.15599996},"labels":[],"label_agreement":null},{"id":"W2110851276","doi":"10.1109/icsm.2006.52","title":"Refactoring Practice: How it is and How it Should be Supported - An Eclipse Case Study","year":2006,"lang":"en","type":"article","venue":"Proceedings/Proceedings - Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":105,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Code refactoring; Eclipse; Computer science; Software engineering; Programming language; Object-oriented programming; Software maintenance; Component (thermodynamics); Development environment; Software development; Software","score_opus":0.08251817860695962,"score_gpt":0.32867406532066457,"score_spread":0.24615588671370495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110851276","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9761342,0.00091616827,0.010622469,0.0027335603,0.000019308854,0.00011500726,0.0000480357,0.00007573707,0.009335558],"genre_scores_gemma":[0.9855993,0.000704602,0.012107745,0.00018702599,0.000011270121,0.000047911377,0.0000775939,0.000034342687,0.0012301375],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.98798764,0.0070688324,0.0006575684,0.00065802195,0.0026272666,0.001000626],"domain_scores_gemma":[0.96972895,0.01984556,0.0024464582,0.002399058,0.0042973547,0.0012826809],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013839758,0.0003889374,0.0003720075,0.0018979512,0.0019289604,0.0025977027,0.0012946461,0.0025445882,0.00049665716],"category_scores_gemma":[0.022449028,0.0003048469,0.0003835747,0.002195811,0.0017654645,0.003798737,0.0015857801,0.0011630374,0.00017904163],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065990014,0.0035086798,0.35924047,0.0011227067,0.000115748444,0.017901607,0.11210779,0.009191633,0.024487145,0.033829663,0.006338495,0.4314962],"study_design_scores_gemma":[0.00043832668,0.003722289,0.4779168,0.0031036905,0.00032952233,0.028597152,0.17478263,0.06656244,0.03957127,0.022571627,0.18189791,0.0005062856],"about_ca_topic_score_codex":0.007593297,"about_ca_topic_score_gemma":0.01652207,"teacher_disagreement_score":0.013839758,"about_ca_system_score_codex":0.003206887,"about_ca_system_score_gemma":0.0016230436,"threshold_uncertainty_score":0.07319248},"labels":[],"label_agreement":null},{"id":"W2110915013","doi":"10.1007/s10664-012-9205-0","title":"Studying the impact of social interactions on software quality","year":2012,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software quality; Software; Quality (philosophy); Source code; Software metric; Data science; Software development; Code review; Software engineering; Data mining","score_opus":0.08293084537093437,"score_gpt":0.3962824852350411,"score_spread":0.31335163986410675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110915013","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9965938,0.00012468782,0.00068678445,0.00027020063,0.0000055506644,0.0000052281,0.000019217458,0.000005888547,0.002288643],"genre_scores_gemma":[0.9993734,0.000055617074,0.00021802832,0.000017977663,0.000009841717,0.0000043816776,0.00001716537,0.0000048436223,0.00029882498],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9956318,0.0027679906,0.0001350225,0.00028081724,0.0008086485,0.0003757444],"domain_scores_gemma":[0.8425017,0.13031778,0.01634014,0.0029258172,0.0041679693,0.0037466197],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041981237,0.00030511495,0.00024752322,0.0013060662,0.0008494249,0.001747814,0.000550904,0.00078796735,0.004425776],"category_scores_gemma":[0.058076896,0.00022837728,0.00041355888,0.0013572081,0.0009903942,0.0022202937,0.0010033867,0.0012496255,0.00031555814],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011258983,0.0029971318,0.8747747,0.00028978902,0.00077292987,0.00047300482,0.005456499,0.009676655,0.0062545342,0.010795258,0.0013962416,0.085987434],"study_design_scores_gemma":[0.00014068307,0.0024199823,0.92610264,0.00008865504,0.00043485826,0.00023379846,0.01089483,0.040515617,0.004218497,0.012130919,0.0027369638,0.00008265624],"about_ca_topic_score_codex":0.0045814714,"about_ca_topic_score_gemma":0.0071206293,"teacher_disagreement_score":0.0045814714,"about_ca_system_score_codex":0.0011823182,"about_ca_system_score_gemma":0.00084831275,"threshold_uncertainty_score":0.022202015},"labels":[],"label_agreement":null},{"id":"W2110987941","doi":"10.1109/wcre.2012.53","title":"Analyzing the impact of antipatterns on change-proneness using fine-grainde source code changes","year":2012,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Java; Computer science; Source code; Class (philosophy); Object-oriented programming; Code (set theory); Programming language; Biology; Artificial intelligence","score_opus":0.04255277538127678,"score_gpt":0.301147974195396,"score_spread":0.25859519881411924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110987941","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9933877,0.0001641544,0.005259517,0.000026475798,0.000009549483,0.00004606263,0.0005407456,0.0001567735,0.00040889287],"genre_scores_gemma":[0.99466175,0.000060265844,0.0041047838,0.000010708672,0.000008726179,0.000024524294,0.0009024678,0.00003723569,0.00018954041],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99547607,0.00078414654,0.0006090781,0.0012292827,0.001660114,0.00024133023],"domain_scores_gemma":[0.88937575,0.067452624,0.024464618,0.010741047,0.0068031717,0.001162766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002589204,0.00040880838,0.0003415315,0.004545296,0.00029059523,0.0006679473,0.00035423454,0.00040403346,0.0007769364],"category_scores_gemma":[0.029302165,0.00023966172,0.00047577586,0.002443059,0.0005682449,0.0010343061,0.00076492375,0.0006273492,0.00014954708],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002713822,0.00012911996,0.922397,0.00017462832,0.00024324039,0.00022569594,0.00046833078,0.0065112542,0.014457971,0.00018127829,0.00020615116,0.054733872],"study_design_scores_gemma":[0.0000056332774,0.00015127401,0.9835548,0.0000119806955,0.00006295661,0.0002146495,0.00013787572,0.010934441,0.0039995844,0.00024803792,0.00066441717,0.000014291179],"about_ca_topic_score_codex":0.0019351984,"about_ca_topic_score_gemma":0.0040092515,"teacher_disagreement_score":0.004545296,"about_ca_system_score_codex":0.00035066117,"about_ca_system_score_gemma":0.00034271844,"threshold_uncertainty_score":0.013693213},"labels":[],"label_agreement":null},{"id":"W2111134789","doi":"10.1145/2591062.2591108","title":"Integrating software project resources using source code identifiers","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Identifier; Source code; Documentation; Android (operating system); KPI-driven code analysis; Class (philosophy); Code review; Software; Open source; Operating system; World Wide Web; Static program analysis; Programming language; Database; Software engineering; Software development","score_opus":0.02994593749957258,"score_gpt":0.2937741858729498,"score_spread":0.2638282483733772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2111134789","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026403219,0.0010341441,0.89071435,0.0014734692,0.00045502622,0.0009031692,0.0062112263,0.036910858,0.035894584],"genre_scores_gemma":[0.115336284,0.0013157314,0.8402268,0.00038389052,0.00018145001,0.00086711656,0.01667318,0.011889075,0.013126439],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9725663,0.007624253,0.00469151,0.0040678084,0.010093798,0.0009564256],"domain_scores_gemma":[0.9063866,0.03560948,0.010021519,0.025831772,0.018767506,0.0033832255],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020030605,0.0021298493,0.001373227,0.027559994,0.001443305,0.00939915,0.002296537,0.0015656919,0.008515907],"category_scores_gemma":[0.10640829,0.0021500345,0.0014304393,0.01714433,0.0012313202,0.017147288,0.012730262,0.0030382988,0.0065213903],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029801365,0.00047302616,0.04638655,0.0013861831,0.00020422669,0.0010193392,0.00847427,0.0030369605,0.0075141876,0.05357494,0.040175445,0.8374569],"study_design_scores_gemma":[0.00013133585,0.0003392493,0.039693996,0.0033629702,0.0004833613,0.0022093605,0.005513862,0.039349146,0.032090347,0.113217294,0.762869,0.0007401489],"about_ca_topic_score_codex":0.0041596736,"about_ca_topic_score_gemma":0.003959879,"teacher_disagreement_score":0.027559994,"about_ca_system_score_codex":0.001522404,"about_ca_system_score_gemma":0.004814694,"threshold_uncertainty_score":0.10593319},"labels":[],"label_agreement":null},{"id":"W2111138279","doi":"10.1109/ms.2005.164","title":"The Art and Science of Software Release Planning","year":2005,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":285,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Enterprise resource planning; Software release life cycle; Product (mathematics); Computer science; New product development; Resource (disambiguation); Process (computing); Software; Software engineering; Process management; Engineering; Software development; Business; Knowledge management; Software quality; Marketing; Operating system","score_opus":0.016706014955432467,"score_gpt":0.2741704459699971,"score_spread":0.2574644310145647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2111138279","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004966874,0.116638884,0.731549,0.034181137,0.0039685876,0.00031731307,0.00096263003,0.0012116506,0.10620399],"genre_scores_gemma":[0.19836712,0.17906056,0.5737477,0.008212718,0.008745916,0.0009491502,0.0014068298,0.00090759713,0.028602444],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98917925,0.0036551687,0.00078548526,0.0017165586,0.004257227,0.0004062927],"domain_scores_gemma":[0.9773267,0.016441325,0.0012172303,0.003108404,0.0015078,0.00039848022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069689844,0.0015699511,0.0014747516,0.003614735,0.0023739485,0.008893364,0.003389872,0.0038865786,0.008936274],"category_scores_gemma":[0.021942452,0.0013698981,0.001466078,0.0052207867,0.015147728,0.009398217,0.0032499724,0.006867345,0.00324755],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006931708,0.000069930065,0.0007192561,0.0011231262,0.000055804972,0.00016094156,0.00076486036,0.025166135,0.00046781247,0.7499298,0.027137196,0.1943359],"study_design_scores_gemma":[0.000033826393,0.000049530096,0.00038782097,0.00051546516,0.00002375168,0.00019144375,0.00026759534,0.014964799,0.0004436254,0.739392,0.24366194,0.000068185116],"about_ca_topic_score_codex":0.009487039,"about_ca_topic_score_gemma":0.004549909,"teacher_disagreement_score":0.009487039,"about_ca_system_score_codex":0.004556734,"about_ca_system_score_gemma":0.0060768165,"threshold_uncertainty_score":0.036855996},"labels":[],"label_agreement":null},{"id":"W2111196330","doi":"10.1109/sedc.1994.475316","title":"A tool for assessing the quality of distributed software designs","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bell (Canada)","funders":"","keywords":"Cyclomatic complexity; Computer science; Measure (data warehouse); Set (abstract data type); Programming language; Transformation (genetics); Software; Quality (philosophy); Theoretical computer science; Petri net; Software quality; Software engineering; Data mining; Software development","score_opus":0.12405510547636843,"score_gpt":0.36580703452519353,"score_spread":0.2417519290488251,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2111196330","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01647195,0.00014618745,0.96601826,0.00009243355,0.000025416706,0.00017516715,0.00056277146,0.013642847,0.0028649522],"genre_scores_gemma":[0.1759363,0.00015341947,0.8196276,0.000062612344,0.000024408246,0.00048994704,0.0013779603,0.0008978932,0.0014298885],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9933963,0.0014815986,0.0006875439,0.00074164296,0.0034753226,0.00021759766],"domain_scores_gemma":[0.98020303,0.011441528,0.0019832724,0.0030525671,0.0029946486,0.00032496825],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040285387,0.00172063,0.00095757167,0.0074947486,0.00067659694,0.0027288445,0.0015699952,0.0011174956,0.0073363124],"category_scores_gemma":[0.029865352,0.0008270497,0.0012297427,0.002727158,0.0010715368,0.0031005575,0.002054857,0.0012196173,0.001416315],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006271856,0.0003113406,0.012138523,0.0012030503,0.00026053228,0.00063112855,0.0012279622,0.11031435,0.038175907,0.13740107,0.014225008,0.6834839],"study_design_scores_gemma":[0.0002610248,0.00080048986,0.005847192,0.00054534554,0.0001775418,0.0013810577,0.0002973345,0.7771993,0.05931257,0.09754765,0.05642936,0.00020113884],"about_ca_topic_score_codex":0.0011995615,"about_ca_topic_score_gemma":0.0009769859,"teacher_disagreement_score":0.0074947486,"about_ca_system_score_codex":0.0011554296,"about_ca_system_score_gemma":0.0013984538,"threshold_uncertainty_score":0.02454245},"labels":[],"label_agreement":null},{"id":"W2111250043","doi":"10.1109/wpc.1998.693281","title":"An analysis framework for understanding layered software architectures","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Nortel (Canada); Bell (Canada)","funders":"","keywords":"Computer science; Software architecture; Layer (electronics); Architecture; Reference architecture; Set (abstract data type); Software engineering; Software; Software architecture description; Multilayered architecture; Computer architecture; Distributed computing; Operating system; Programming language","score_opus":0.06398928410912212,"score_gpt":0.30682040875364586,"score_spread":0.24283112464452372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2111250043","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031612394,0.00025655265,0.9927618,0.0005482153,0.000012924777,0.00011416224,0.00007106737,0.0003717327,0.0027022164],"genre_scores_gemma":[0.08670007,0.00035875404,0.91084933,0.00012760954,0.000034545435,0.00044771444,0.0003310805,0.00018948925,0.0009614037],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9945328,0.002244544,0.0004834346,0.000647497,0.0017404999,0.00035109857],"domain_scores_gemma":[0.98816174,0.0068468386,0.0012837223,0.0015251552,0.0019222919,0.00026008804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010062515,0.0023723193,0.0009699076,0.00781941,0.0027396635,0.006731742,0.0028944074,0.0022323288,0.0043056225],"category_scores_gemma":[0.020592108,0.0013643763,0.0024774924,0.0037222693,0.0065639415,0.017360479,0.003321318,0.0035829945,0.0009905013],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000029428242,0.00006378878,0.0023329908,0.00029043498,0.000082422084,0.00030628062,0.0053256294,0.03135337,0.0036714824,0.9003541,0.0018225196,0.054367546],"study_design_scores_gemma":[0.000019037378,0.000044713808,0.00083916105,0.0002497749,0.000049193444,0.00022896077,0.0015605008,0.13518743,0.0022838244,0.8329918,0.026480153,0.000065516455],"about_ca_topic_score_codex":0.010579525,"about_ca_topic_score_gemma":0.005271441,"teacher_disagreement_score":0.010579525,"about_ca_system_score_codex":0.0040002535,"about_ca_system_score_gemma":0.0040376387,"threshold_uncertainty_score":0.05321628},"labels":[],"label_agreement":null},{"id":"W2111555385","doi":"10.1109/tse.2010.70","title":"Solving the Class Responsibility Assignment Problem in Object-Oriented Analysis with Multi-Objective Genetic Algorithms","year":2010,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":128,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Heuristics; Class (philosophy); Genetic algorithm; Cohesion (chemistry); Domain (mathematical analysis); Context (archaeology); Class diagram; Object-oriented programming; Algorithm; Machine learning; Artificial intelligence; Programming language; Mathematics; Unified Modeling Language; Software","score_opus":0.009970145975078455,"score_gpt":0.24159022935009608,"score_spread":0.23162008337501763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2111555385","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042622704,0.00018163334,0.9551,0.00018836801,0.000018242397,0.00006618105,0.000013570233,0.00017217145,0.001637144],"genre_scores_gemma":[0.36823687,0.00020813692,0.62985575,0.000112188,0.00002221478,0.00019985947,0.00004391691,0.00008325003,0.0012377573],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99920744,0.00043675138,0.000030345074,0.00010041872,0.00015422805,0.00007080374],"domain_scores_gemma":[0.9982084,0.0013428895,0.00018722491,0.000059456648,0.00015051565,0.000051500065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002317964,0.0009845312,0.00085134694,0.0013040017,0.0008032504,0.0012034951,0.0009739937,0.0014586229,0.00094693084],"category_scores_gemma":[0.004436675,0.0005998877,0.0007397118,0.00089269894,0.0011040715,0.0008876045,0.00087495986,0.0009506788,0.0001380574],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019604222,0.00003502528,0.000584025,0.000042489086,0.000029581593,0.00004022149,0.00008733598,0.96635544,0.00065110705,0.0062113344,0.00021386963,0.025729919],"study_design_scores_gemma":[0.000011942895,0.00001733558,0.000096227115,0.000009888722,0.000012356176,0.000014109717,0.000030335508,0.9922908,0.00046882793,0.0067099305,0.00033346887,0.000004800983],"about_ca_topic_score_codex":0.0059367674,"about_ca_topic_score_gemma":0.0056967414,"teacher_disagreement_score":0.0059367674,"about_ca_system_score_codex":0.0012263235,"about_ca_system_score_gemma":0.0018508389,"threshold_uncertainty_score":0.0122587085},"labels":[],"label_agreement":null},{"id":"W2111937777","doi":"10.1109/isa.2008.104","title":"Catalog of Metrics for Assessing Security Risks of Software throughout the Software Development Life Cycle","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Systems development life cycle; Software security assurance; Computer science; Software development process; Software development; Software peer review; Software engineering; Software metric; Security testing; Security engineering; Security bug; Computer security; Security information and event management; Risk analysis (engineering); Software construction; Software; Security service; Information security; Cloud computing security; Business","score_opus":0.09007437702113673,"score_gpt":0.35672487196033914,"score_spread":0.2666504949392024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2111937777","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08366268,0.01253505,0.831686,0.0021153379,0.0004697119,0.0042531216,0.012969724,0.013862866,0.038445592],"genre_scores_gemma":[0.20807357,0.0063039055,0.7625153,0.00018129079,0.000108170585,0.0031898692,0.014810591,0.00091175735,0.0039054852],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9757264,0.004537475,0.0052106366,0.0011044255,0.01297084,0.00045028984],"domain_scores_gemma":[0.92444634,0.019930054,0.012770985,0.014559782,0.02661498,0.0016779276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014858971,0.0023827925,0.0017417654,0.023676978,0.0015111838,0.004825355,0.0018804112,0.0013264953,0.0019655984],"category_scores_gemma":[0.056948796,0.0009047756,0.0014617243,0.010933944,0.0007147364,0.005623667,0.0020385543,0.0016030304,0.0012351673],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029375183,0.0008124985,0.07904409,0.004564896,0.0005987767,0.0002704543,0.001161683,0.024804188,0.013107036,0.061288282,0.021110008,0.79294443],"study_design_scores_gemma":[0.00031889288,0.004903441,0.1813677,0.008941463,0.001889877,0.0029119719,0.0018394446,0.22155407,0.07259074,0.101983294,0.40057915,0.0011200173],"about_ca_topic_score_codex":0.0039369906,"about_ca_topic_score_gemma":0.004817848,"teacher_disagreement_score":0.023676978,"about_ca_system_score_codex":0.0027418362,"about_ca_system_score_gemma":0.0057629864,"threshold_uncertainty_score":0.078582704},"labels":[],"label_agreement":null},{"id":"W2112332360","doi":"10.1109/ccece.2004.1345050","title":"Knowledge representation and processing in intelligent software measurement system (ISMS)","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Knowledge base; Knowledge-based systems; Software; Plan (archaeology); Component (thermodynamics); Expert system; Open Knowledge Base Connectivity; Human–computer interaction; Software engineering; Knowledge management; Artificial intelligence; Personal knowledge management; Operating system","score_opus":0.058781294408131665,"score_gpt":0.30309730452134787,"score_spread":0.2443160101132162,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112332360","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012149962,0.0008425175,0.9758814,0.0010553517,0.000046351546,0.00019137135,0.0002880187,0.0011037053,0.008441402],"genre_scores_gemma":[0.18279824,0.0011841701,0.8101432,0.00033439585,0.000070441034,0.00041717503,0.0011315929,0.0000789914,0.0038418777],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99804425,0.00061514723,0.0002870681,0.0003587643,0.00055026525,0.00014452347],"domain_scores_gemma":[0.9983277,0.0007904401,0.00016082705,0.00036547662,0.00029073344,0.00006493639],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020238266,0.00044162196,0.00062375795,0.00248235,0.0009754397,0.0039196797,0.0021130976,0.001483524,0.0024022073],"category_scores_gemma":[0.005947358,0.0005661063,0.0010424461,0.002891952,0.0017253442,0.004973,0.001760359,0.0012516213,0.0009858574],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021045395,0.00031213555,0.002606302,0.000664845,0.00013951315,0.00147555,0.0022754034,0.14954658,0.0071757296,0.35428134,0.011727615,0.46958452],"study_design_scores_gemma":[0.00005880356,0.00009433128,0.0013267957,0.00027387106,0.00014717226,0.00051519595,0.0006217238,0.5870188,0.009529772,0.3519216,0.048415206,0.00007672408],"about_ca_topic_score_codex":0.011309408,"about_ca_topic_score_gemma":0.0092116725,"teacher_disagreement_score":0.011309408,"about_ca_system_score_codex":0.0019486544,"about_ca_system_score_gemma":0.0023192589,"threshold_uncertainty_score":0.022487164},"labels":[],"label_agreement":null},{"id":"W2112433871","doi":"10.1145/1134285.1134333","title":"On the success of empirical studies in the international conference on software engineering","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":130,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Empirical research; Soundness; Computer science; Point (geometry); Software engineering; Software quality; Software; Quality (philosophy); Management science; Data science; Software development; Engineering; Programming language; Mathematics; Epistemology","score_opus":0.09787210732289231,"score_gpt":0.3613300767234555,"score_spread":0.2634579694005632,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112433871","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4290594,0.101516336,0.082292385,0.1148473,0.006443825,0.00257252,0.0015009112,0.00041147013,0.26135588],"genre_scores_gemma":[0.9655107,0.007835664,0.013855469,0.0065248585,0.0012288683,0.0012155301,0.00039600048,0.00018910448,0.0032438023],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.44776204,0.3929068,0.027620248,0.01630945,0.111870736,0.003530661],"domain_scores_gemma":[0.047666505,0.8069263,0.03653058,0.03316179,0.07370698,0.0020078607],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.39018053,0.0015133201,0.0018920611,0.010343248,0.0053917817,0.015740331,0.00351994,0.004772105,0.0067537106],"category_scores_gemma":[0.8209343,0.00091682415,0.0011208443,0.014325107,0.017056195,0.015132803,0.007184146,0.0056212037,0.0013398681],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029920815,0.00206549,0.14409736,0.010503686,0.0027817425,0.00040828256,0.046686735,0.00434444,0.0023506605,0.31450713,0.046110343,0.42315203],"study_design_scores_gemma":[0.0014252956,0.006983007,0.34701923,0.035851177,0.0022814865,0.0012407575,0.0571233,0.014342168,0.018074356,0.22106166,0.29352516,0.0010724603],"about_ca_topic_score_codex":0.0026262014,"about_ca_topic_score_gemma":0.00294695,"teacher_disagreement_score":0.98965675,"about_ca_system_score_codex":0.011499126,"about_ca_system_score_gemma":0.00718197,"threshold_uncertainty_score":0.752016},"labels":[],"label_agreement":null},{"id":"W2112533109","doi":"10.1145/1101908.1101919","title":"UMLDiff","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":369,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Correctness; Software evolution; Java; Robustness (evolution); Object-oriented design; Software system; Software; Software engineering; Software maintenance; Software design; Software development; Structural pattern; Programming language; Software construction","score_opus":0.014352778789238047,"score_gpt":0.26290284714321815,"score_spread":0.2485500683539801,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112533109","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015803447,0.00077232864,0.8831563,0.00061651587,0.00034044954,0.0005625561,0.009685855,0.08719399,0.0160916],"genre_scores_gemma":[0.016170235,0.0009004426,0.921221,0.0005932692,0.00013399581,0.0009484226,0.032692418,0.015052635,0.012287544],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99220467,0.0019666972,0.0015021262,0.0011312516,0.0028077974,0.00038738974],"domain_scores_gemma":[0.9875124,0.004819369,0.0010416355,0.0036406757,0.0026007781,0.00038512246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007798735,0.0022147663,0.0016007044,0.005996014,0.0019526455,0.0058635175,0.005480699,0.002618357,0.029965376],"category_scores_gemma":[0.022046946,0.0028661625,0.0037712154,0.003404647,0.001048214,0.0073906346,0.0056236964,0.0037593425,0.025894247],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041825775,0.0002233681,0.0038461443,0.001746366,0.00021962666,0.00043715563,0.0010670843,0.013508358,0.0053982367,0.18940012,0.182444,0.60129124],"study_design_scores_gemma":[0.00010460441,0.00005603745,0.0004951406,0.00026638,0.0000575985,0.000589365,0.000102103135,0.027590442,0.0062859785,0.07296108,0.89139056,0.00010069422],"about_ca_topic_score_codex":0.004206015,"about_ca_topic_score_gemma":0.00541132,"teacher_disagreement_score":0.029965376,"about_ca_system_score_codex":0.0019600387,"about_ca_system_score_gemma":0.0035473462,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2112584502","doi":"10.1109/wcre.2003.1287241","title":"Reverse engineering the process of small novice software teams","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Team software process; Personal software process; Software project management; Software development process; Software Engineering Process Group; Computer science; Software development; Software engineering; Process (computing); Task (project management); Adaptation (eye); Competence (human resources); Engineering management; Process management; Software; Engineering; Systems engineering; Software construction","score_opus":0.01006665361565912,"score_gpt":0.23715432602046305,"score_spread":0.22708767240480393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112584502","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6813733,0.00064916804,0.296848,0.00095796137,0.00009022261,0.00096115225,0.00019056915,0.0011787559,0.017750924],"genre_scores_gemma":[0.79788023,0.0004162539,0.17870134,0.00018761192,0.00004153896,0.0002810113,0.00048220722,0.0002974611,0.021712337],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9924663,0.003909055,0.00031942624,0.0008434087,0.0020432216,0.00041851582],"domain_scores_gemma":[0.95341396,0.023456953,0.003266668,0.01218111,0.0060931896,0.0015880867],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008376064,0.00046960032,0.00042622272,0.0019959176,0.0022285169,0.0030034704,0.0013205914,0.0010195096,0.0042163325],"category_scores_gemma":[0.03972455,0.00048271235,0.00048640327,0.00088599144,0.0021332987,0.0032567035,0.0032453989,0.0010928405,0.0012620938],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007297994,0.0011517301,0.0749472,0.00073104526,0.00012007988,0.0027084416,0.045194138,0.02352307,0.027172977,0.037544254,0.0051660147,0.7810112],"study_design_scores_gemma":[0.00036438135,0.0038555427,0.10481868,0.001143799,0.00027652786,0.0066113095,0.051991675,0.25159973,0.12323398,0.17792122,0.27764478,0.0005384427],"about_ca_topic_score_codex":0.0037273068,"about_ca_topic_score_gemma":0.0067994315,"teacher_disagreement_score":0.008376064,"about_ca_system_score_codex":0.0011048943,"about_ca_system_score_gemma":0.0030649463,"threshold_uncertainty_score":0.044297397},"labels":[],"label_agreement":null},{"id":"W2112768789","doi":"10.1109/compsac.2012.58","title":"Doppel-Code: A Clone Visualization Tool for Prioritizing Global and Local Clone Impacts","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"clone (Java method); Visualization; Computer science; Software engineering; Software maintenance; Program comprehension; Source code; Software visualization; Code (set theory); Software; Software development; Programming language; Software system; Data mining; Software construction; Biology","score_opus":0.021654686214798335,"score_gpt":0.3248974695766934,"score_spread":0.3032427833618951,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112768789","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09840786,0.0017381717,0.6299899,0.00075899035,0.00018142053,0.0005459035,0.014153261,0.2469495,0.0072749797],"genre_scores_gemma":[0.2852326,0.0008322365,0.68379784,0.00020867275,0.000065631495,0.0005328292,0.014125526,0.011269618,0.003935144],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983311,0.0002595392,0.0001991401,0.00029907786,0.0008175882,0.00009361312],"domain_scores_gemma":[0.98610675,0.008036518,0.0016412039,0.0011796345,0.0025672463,0.0004686573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024419704,0.0017483068,0.0010295769,0.009192677,0.0008710714,0.0022600116,0.0017644631,0.001200679,0.0051941164],"category_scores_gemma":[0.018033236,0.00070487877,0.00066807616,0.004397164,0.00043668848,0.0039582415,0.0027303488,0.0013156878,0.0011369934],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010271426,0.0002898641,0.057233933,0.0019954098,0.0003322896,0.0010833449,0.0035009277,0.012090913,0.032649755,0.009345028,0.07828942,0.8021619],"study_design_scores_gemma":[0.0005787075,0.000784022,0.059956867,0.0008964164,0.00048421856,0.003274753,0.0020704074,0.56741536,0.13872606,0.030978959,0.19418219,0.0006520364],"about_ca_topic_score_codex":0.0042848354,"about_ca_topic_score_gemma":0.0074810837,"teacher_disagreement_score":0.009192677,"about_ca_system_score_codex":0.00056305807,"about_ca_system_score_gemma":0.0014460724,"threshold_uncertainty_score":0.017376065},"labels":[],"label_agreement":null},{"id":"W2112857338","doi":"10.1145/1866272.1866279","title":"Specifying overlaps of heterogeneous models for global consistency checking","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Consistency (knowledge bases); Computer science; Metamodeling; Sketch; Consistency model; Set (abstract data type); Semantics (computer science); Theoretical computer science; Model checking; Formal specification; Sequential consistency; Programming language; Artificial intelligence; Algorithm","score_opus":0.03785611253070094,"score_gpt":0.2901792170803722,"score_spread":0.2523231045496713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112857338","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01069478,0.000050860865,0.9871022,0.000111542475,0.000010202824,0.0001287339,0.000058251408,0.0005675241,0.0012758746],"genre_scores_gemma":[0.27389258,0.00010914074,0.7233453,0.00015161018,0.000041640702,0.000621626,0.00037177047,0.0006413726,0.0008249945],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97103184,0.012271245,0.0027634155,0.0028616192,0.008702148,0.0023696763],"domain_scores_gemma":[0.96149874,0.022346724,0.0026865734,0.010072567,0.002694781,0.00070066616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020117087,0.0017693951,0.0020324953,0.002897574,0.0026181245,0.004038601,0.0031919472,0.0022736846,0.003201803],"category_scores_gemma":[0.03333196,0.0021655313,0.0037483114,0.0023028487,0.005337519,0.011404879,0.011912776,0.005395427,0.00042129666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029726545,0.00018780715,0.0060846303,0.00033775176,0.00022184283,0.0014702802,0.00482576,0.10454888,0.011436651,0.81947184,0.002029544,0.04908773],"study_design_scores_gemma":[0.00012939943,0.00017990329,0.00084990467,0.00047582577,0.00033546792,0.00073912303,0.0017399342,0.44504285,0.029971989,0.49338895,0.026998965,0.000147661],"about_ca_topic_score_codex":0.0043036197,"about_ca_topic_score_gemma":0.006797229,"teacher_disagreement_score":0.020117087,"about_ca_system_score_codex":0.0024955594,"about_ca_system_score_gemma":0.003880117,"threshold_uncertainty_score":0.106390595},"labels":[],"label_agreement":null},{"id":"W2112890525","doi":"10.1145/1985404.1985425","title":"VisCad","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Visualization; clone (Java method); Computer science; Cloning (programming); Data visualization; Data mining; Programming language; Biology","score_opus":0.04177068418181956,"score_gpt":0.25236575569313063,"score_spread":0.21059507151131107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112890525","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010574783,0.00086576294,0.32863283,0.00071370264,0.00032867768,0.00022843428,0.019543065,0.6018724,0.03724039],"genre_scores_gemma":[0.14033018,0.0015466627,0.6285288,0.0013798835,0.00018627146,0.0010215867,0.052525006,0.10611355,0.06836809],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99906296,0.00011338606,0.00006607584,0.00018004837,0.0004921594,0.00008536396],"domain_scores_gemma":[0.9971307,0.0011700178,0.00017051645,0.000736609,0.0006881939,0.00010394981],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014199818,0.001562255,0.00090175384,0.004010735,0.00068128144,0.0022716518,0.001992335,0.0011891191,0.06907241],"category_scores_gemma":[0.00713382,0.0008037653,0.0008981361,0.0017595782,0.0003375293,0.003432374,0.002854822,0.0016135212,0.02160449],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008983856,0.00012899122,0.0061145183,0.0011797068,0.00011584717,0.0004908795,0.0009170404,0.0031122363,0.02231377,0.023719925,0.43765727,0.50335145],"study_design_scores_gemma":[0.00029651143,0.00022041476,0.007597481,0.00047227624,0.00011848161,0.0021774187,0.00038470887,0.07304869,0.062860005,0.028795738,0.8237659,0.000262299],"about_ca_topic_score_codex":0.0016499747,"about_ca_topic_score_gemma":0.0025472487,"teacher_disagreement_score":0.06907241,"about_ca_system_score_codex":0.0006370203,"about_ca_system_score_gemma":0.0007201226,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2112947311","doi":"10.1109/wse.2012.6320536","title":"Normalizing object-oriented class styles in JavaScript","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"JavaScript; Unobtrusive JavaScript; Computer science; Programming language; Program comprehension; Maintainability; Object-oriented programming; Web application; Flexibility (engineering); Software engineering; Class (philosophy); Rich Internet application; Representation (politics); World Wide Web; Artificial intelligence; Software; Software system","score_opus":0.019201196509376732,"score_gpt":0.2632359801479525,"score_spread":0.24403478363857578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112947311","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013362393,0.00017996007,0.97011596,0.0001505913,0.00016020946,0.00011647652,0.00026942085,0.0112792095,0.0043657375],"genre_scores_gemma":[0.10795924,0.00067289453,0.8680364,0.00045060815,0.000117679614,0.00031749444,0.00103734,0.014636118,0.006772246],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9965702,0.00056425895,0.00043069915,0.000689708,0.0015736856,0.00017154001],"domain_scores_gemma":[0.9916351,0.0020709164,0.000914327,0.0036464958,0.0014967636,0.0002363328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023703924,0.00067817396,0.00068330206,0.0011634085,0.0006717456,0.0035791392,0.0017215811,0.0007790597,0.0018409429],"category_scores_gemma":[0.012039699,0.0009001915,0.0007251063,0.0016341234,0.0016426961,0.0038128945,0.0019848808,0.0030922778,0.0027658343],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040739868,0.00031297765,0.0075352876,0.0007186424,0.00008835798,0.00077724067,0.005729118,0.01957086,0.1770539,0.25775453,0.017291905,0.51275986],"study_design_scores_gemma":[0.00014104269,0.00009463506,0.0038754274,0.00041346057,0.00016233171,0.0015246606,0.00049721176,0.11023003,0.25628042,0.15934435,0.46714726,0.00028924434],"about_ca_topic_score_codex":0.0017563441,"about_ca_topic_score_gemma":0.0020383738,"teacher_disagreement_score":0.0035791392,"about_ca_system_score_codex":0.00084512663,"about_ca_system_score_gemma":0.001601045,"threshold_uncertainty_score":0.01253593},"labels":[],"label_agreement":null},{"id":"W2113200858","doi":"10.1109/hcc.2002.1046339","title":"Quantifying developer experiences via heuristic and psychometric evaluation","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Learnability; Usability; Computer science; Heuristic evaluation; Java; Relation (database); Heuristic; Visibility; Conceptual model; Software engineering; Human–computer interaction; Programming language; Artificial intelligence; Data mining; Database","score_opus":0.10075377670311164,"score_gpt":0.3549333827517294,"score_spread":0.2541796060486178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113200858","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97814375,0.00014588953,0.015729865,0.00006349308,0.00001572723,0.0008256602,0.00018839077,0.000114221846,0.004772978],"genre_scores_gemma":[0.9777827,0.0001363736,0.019611368,0.00003639309,0.00001461539,0.0013957247,0.0005057125,0.000028819906,0.00048830605],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.94356406,0.03369067,0.007003912,0.0019835997,0.012482013,0.0012757954],"domain_scores_gemma":[0.6476889,0.25580892,0.02096601,0.01953811,0.053575933,0.0024221353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.055493653,0.0008014112,0.0007024977,0.007621213,0.00079412054,0.0020532815,0.00076413795,0.0007893762,0.00085956213],"category_scores_gemma":[0.22750801,0.00036828071,0.0007019617,0.004706656,0.001305363,0.0016736652,0.0026546447,0.00061915815,0.00037472238],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008618438,0.0024428575,0.6799057,0.0006339436,0.00024824913,0.00025118524,0.027219512,0.004946935,0.0035546147,0.0015795657,0.001168233,0.27718726],"study_design_scores_gemma":[0.00045067494,0.01145112,0.8626573,0.00056255073,0.00032538894,0.0009559354,0.045388296,0.042747512,0.015872836,0.0072193616,0.011975541,0.0003935223],"about_ca_topic_score_codex":0.0010432444,"about_ca_topic_score_gemma":0.0014688603,"teacher_disagreement_score":0.055493653,"about_ca_system_score_codex":0.0012013036,"about_ca_system_score_gemma":0.001155887,"threshold_uncertainty_score":0.293482},"labels":[],"label_agreement":null},{"id":"W2113322762","doi":"10.1109/qsic.2009.47","title":"A Bayesian Approach for the Detection of Code and Design Smells","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":223,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"","keywords":"Code smell; Computer science; Context (archaeology); Probabilistic logic; Process (computing); Code (set theory); Machine learning; Artificial intelligence; Bayesian probability; Data mining; Quality (philosophy); Source code; Statistical model; Software quality; Programming language; Software development; Software","score_opus":0.03149184009471971,"score_gpt":0.2614424345313111,"score_spread":0.22995059443659138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113322762","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010042392,0.00013933067,0.98754483,0.00036414963,0.0000129163145,0.000055535456,0.00015214169,0.00059366826,0.0010950807],"genre_scores_gemma":[0.35068777,0.0003200837,0.6439387,0.00033277107,0.00010086982,0.00037441784,0.00071753154,0.00028690996,0.0032409837],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99340713,0.0022656978,0.00035906644,0.0013822942,0.0022357223,0.00035010115],"domain_scores_gemma":[0.9731496,0.018812839,0.0024777886,0.0020336611,0.003073111,0.00045304187],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007769622,0.0012290856,0.0017325191,0.0059745335,0.0013818194,0.0030649768,0.0042756153,0.003149654,0.0031948222],"category_scores_gemma":[0.04748438,0.0017357648,0.0021572278,0.00267462,0.00200986,0.004582217,0.0023944802,0.0031075913,0.0009852396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041780554,0.00048615722,0.01920076,0.0003591395,0.00030983594,0.00041788843,0.0010306303,0.602597,0.007497902,0.12018515,0.0059602195,0.24153751],"study_design_scores_gemma":[0.000020392526,0.000030333298,0.0014455739,0.000033225708,0.000023225208,0.00014062431,0.000028309301,0.9477156,0.00086723216,0.048513982,0.0011432115,0.000038247144],"about_ca_topic_score_codex":0.017160583,"about_ca_topic_score_gemma":0.020584455,"teacher_disagreement_score":0.017160583,"about_ca_system_score_codex":0.002604098,"about_ca_system_score_gemma":0.0027614736,"threshold_uncertainty_score":0.04109019},"labels":[],"label_agreement":null},{"id":"W2113351233","doi":"10.1145/2000791.2000794","title":"Reducing the effort of bug report triage","year":2011,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":281,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Process (computing); Variety (cybernetics); Key (lock); Recommender system; Software engineering; Software; Software development; Triage; Data science; World Wide Web; Artificial intelligence; Computer security","score_opus":0.13501545347678479,"score_gpt":0.32812057111168713,"score_spread":0.19310511763490235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113351233","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31704542,0.0022642368,0.5776415,0.002847219,0.00061548594,0.0019101911,0.0009615791,0.08520993,0.011504425],"genre_scores_gemma":[0.47875762,0.00052836287,0.5046675,0.000497588,0.00025787434,0.00064379716,0.002144687,0.0030150781,0.009487489],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.97731394,0.008962637,0.0019120545,0.0035357364,0.007493264,0.0007823518],"domain_scores_gemma":[0.8559619,0.065880224,0.013762273,0.0408127,0.020237021,0.0033459803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015081359,0.002487269,0.002079073,0.005286417,0.0014714202,0.0032664011,0.004399469,0.0023507446,0.004200506],"category_scores_gemma":[0.12860788,0.0015993713,0.0013488401,0.0026194432,0.00075777806,0.005054185,0.0040662717,0.0027902394,0.0055732327],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075109396,0.0012075446,0.03174909,0.0004650582,0.00016539785,0.00033806096,0.0025370428,0.015125246,0.023525584,0.001578991,0.020393774,0.902163],"study_design_scores_gemma":[0.0006759297,0.003977825,0.07669143,0.00049313763,0.0007387581,0.0024754978,0.0037484618,0.7444115,0.06900931,0.0067415196,0.09050238,0.00053431396],"about_ca_topic_score_codex":0.009170181,"about_ca_topic_score_gemma":0.010913923,"teacher_disagreement_score":0.015081359,"about_ca_system_score_codex":0.0013726144,"about_ca_system_score_gemma":0.00374247,"threshold_uncertainty_score":0.07975882},"labels":[],"label_agreement":null},{"id":"W2113476536","doi":"10.1109/icsm.2005.42","title":"Dynamic feature traces: finding features in unfamiliar code","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":126,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Heuristics; Feature (linguistics); Relevance (law); Test suite; TRACE (psycholinguistics); Ranking (information retrieval); Code (set theory); Suite; Source code; Data mining; Binary code; Quality (philosophy); Artificial intelligence; Binary number; Machine learning; Information retrieval; Test case; Programming language","score_opus":0.01121693938005001,"score_gpt":0.27934726706974516,"score_spread":0.26813032768969514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113476536","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2865945,0.00026756115,0.6966803,0.0002914112,0.000037595622,0.00031825592,0.0005215234,0.012617438,0.002671332],"genre_scores_gemma":[0.7148464,0.00009910735,0.28033578,0.00005885135,0.000027751888,0.0001307905,0.000746645,0.0008213698,0.0029333122],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975872,0.000645128,0.00011432702,0.00047791132,0.0009860508,0.00018934124],"domain_scores_gemma":[0.98393494,0.008977678,0.0027797192,0.0021626786,0.0016997715,0.00044509408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016201548,0.0010042733,0.00064127793,0.003958873,0.0008693136,0.0014786883,0.001800631,0.00094895036,0.0025248744],"category_scores_gemma":[0.019622345,0.00045003556,0.00044260817,0.0021283794,0.00076039997,0.0032467535,0.0018189836,0.00083596923,0.0008852467],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066591625,0.00040608758,0.057261467,0.00048479965,0.00010541767,0.0010005666,0.003463924,0.011019609,0.039831072,0.0076493258,0.0043884567,0.8737234],"study_design_scores_gemma":[0.00020755996,0.0016587949,0.07171782,0.0002873315,0.00032418227,0.003681626,0.0046432433,0.6532717,0.16401596,0.056984536,0.042825665,0.00038151952],"about_ca_topic_score_codex":0.003755847,"about_ca_topic_score_gemma":0.007254102,"teacher_disagreement_score":0.003958873,"about_ca_system_score_codex":0.00061724416,"about_ca_system_score_gemma":0.0012885001,"threshold_uncertainty_score":0.0085683465},"labels":[],"label_agreement":null},{"id":"W2113509293","doi":"10.1109/icsm.2008.4658065","title":"Understanding the rationale for updating a function’s comment","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Function (biology); Code (set theory); Compiler; Source code; Data mining; Cluster analysis; Software; Machine learning; Programming language","score_opus":0.16180401284416618,"score_gpt":0.2878039930766869,"score_spread":0.12599998023252074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113509293","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94084847,0.0003855968,0.05066193,0.0017099759,0.000065967804,0.00028219784,0.0013195624,0.00048404327,0.0042422162],"genre_scores_gemma":[0.9767462,0.00013807944,0.020771312,0.00012789114,0.00004191192,0.00006188264,0.0010283897,0.00012324932,0.00096108305],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.986917,0.0060952026,0.0013955164,0.0013989423,0.0036277778,0.0005655978],"domain_scores_gemma":[0.57809067,0.3160436,0.05331918,0.02058415,0.030522041,0.0014402904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022309493,0.0005791765,0.00040192888,0.0030169608,0.00097359205,0.0019328983,0.0010289811,0.001907851,0.0019533825],"category_scores_gemma":[0.22707868,0.00048702964,0.0006162484,0.0017844153,0.0013424053,0.004291049,0.0007479292,0.0015318751,0.0006277867],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005257144,0.00016389783,0.8521865,0.000351991,0.0001486269,0.0010999619,0.0033833848,0.012192893,0.0052953665,0.0055828453,0.0029124904,0.11615628],"study_design_scores_gemma":[0.000097949625,0.00039268882,0.7684799,0.0006571141,0.00028667843,0.0037772926,0.006678225,0.15720569,0.018467302,0.015571938,0.028039157,0.00034611003],"about_ca_topic_score_codex":0.0074722837,"about_ca_topic_score_gemma":0.009706556,"teacher_disagreement_score":0.022309493,"about_ca_system_score_codex":0.0013541959,"about_ca_system_score_gemma":0.001504852,"threshold_uncertainty_score":0.11798531},"labels":[],"label_agreement":null},{"id":"W2113642562","doi":"10.1109/msr.2007.28","title":"Release Pattern Discovery via Partitioning: Methodology and Case Study","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science","score_opus":0.07319481911679364,"score_gpt":0.35398485389792206,"score_spread":0.2807900347811284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113642562","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54772776,0.00041339983,0.44469115,0.00033730446,0.000024776753,0.002606405,0.0006900113,0.00041458267,0.0030947027],"genre_scores_gemma":[0.4574421,0.0003629258,0.53747547,0.0000812799,0.00002484416,0.001914137,0.0010384849,0.00011634373,0.0015443681],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9899969,0.004863948,0.00081124273,0.0019653046,0.0018305241,0.0005319955],"domain_scores_gemma":[0.967046,0.023186577,0.0034137266,0.0032304365,0.0025357823,0.0005874016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006119906,0.00075486087,0.0006837907,0.005128546,0.0014985515,0.0017566377,0.002352962,0.0018688687,0.001124496],"category_scores_gemma":[0.019951763,0.0007244732,0.0009707526,0.004628745,0.0016824283,0.001676647,0.0023477727,0.0008899858,0.00033219325],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009900709,0.0053616674,0.20334941,0.0029002114,0.00039654365,0.013753621,0.04655795,0.03731188,0.07211463,0.026618421,0.004524749,0.58612096],"study_design_scores_gemma":[0.0007327706,0.0032755046,0.13354631,0.00088877406,0.0009017866,0.025252461,0.04211619,0.5468401,0.15965295,0.029457979,0.05673603,0.0005990843],"about_ca_topic_score_codex":0.003953777,"about_ca_topic_score_gemma":0.0057024183,"teacher_disagreement_score":0.006119906,"about_ca_system_score_codex":0.0010728391,"about_ca_system_score_gemma":0.0016811513,"threshold_uncertainty_score":0.03236556},"labels":[],"label_agreement":null},{"id":"W2113772893","doi":"10.1016/j.scico.2013.03.015","title":"Using heuristics to estimate an appropriate number of latent topics in source code analysis","year":2013,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Latent Dirichlet allocation; Source code; Topic model; Heuristics; Set (abstract data type); Locality; Cluster analysis; Code review; Code (set theory); KPI-driven code analysis; Software; Programming language; Theoretical computer science; Information retrieval; Static program analysis; Software development; Artificial intelligence","score_opus":0.03221342878942683,"score_gpt":0.34217966633226143,"score_spread":0.3099662375428346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113772893","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059456665,0.00047810926,0.9372499,0.00025952808,0.00004116278,0.00015137186,0.00031780312,0.0016218156,0.00042370253],"genre_scores_gemma":[0.42958987,0.0002581458,0.5659839,0.00020209132,0.00008717525,0.0005339192,0.0022925832,0.0003777079,0.00067459757],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950288,0.0026040738,0.0003225575,0.001087149,0.00053265097,0.00042482925],"domain_scores_gemma":[0.9615428,0.03343609,0.0011102771,0.0014549514,0.0019320542,0.0005239409],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006223495,0.0015416779,0.0019996576,0.005413289,0.0017197939,0.0031746514,0.0022800516,0.0027208468,0.0016544014],"category_scores_gemma":[0.039250858,0.0014480333,0.002053682,0.0038238624,0.0013187672,0.0043895845,0.0019030201,0.003874375,0.0007225],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021542604,0.0012528852,0.035645645,0.0008120415,0.00082866347,0.00029378434,0.002005057,0.2527653,0.019818718,0.026195155,0.014661975,0.6435665],"study_design_scores_gemma":[0.00014046601,0.000055204095,0.0026779496,0.00004209265,0.00011401427,0.00006486457,0.00018184258,0.9671396,0.0028378016,0.02600445,0.0007042755,0.00003749715],"about_ca_topic_score_codex":0.011493518,"about_ca_topic_score_gemma":0.017972546,"teacher_disagreement_score":0.011493518,"about_ca_system_score_codex":0.0018337518,"about_ca_system_score_gemma":0.0033270686,"threshold_uncertainty_score":0.032913387},"labels":[],"label_agreement":null},{"id":"W2113938691","doi":"10.1109/scam.2015.7335412","title":"On the comprehension of code clone visualizations: A controlled study using eye tracking","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Program comprehension; Computer science; Eye tracking; clone (Java method); Code (set theory); Comprehension; Programming language; Tracking (education); Artificial intelligence; Human–computer interaction; Software; Software system; Psychology; Biology; Genetics","score_opus":0.11487371856370028,"score_gpt":0.38297117724606106,"score_spread":0.2680974586823608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113938691","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99838495,0.000050490115,0.0006122483,0.00003072083,0.0000061962296,0.00039584952,0.000056492456,0.000019173247,0.00044397736],"genre_scores_gemma":[0.9935715,0.000116688316,0.0029912076,0.00013768532,0.000017986817,0.0018703527,0.00013183522,0.000023514645,0.001139277],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9974335,0.0010040036,0.00020225887,0.00069702626,0.0003586237,0.00030464257],"domain_scores_gemma":[0.9726183,0.019772641,0.0027809215,0.0012344691,0.0024635906,0.0011300201],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00491933,0.0008727941,0.00073246076,0.00092568254,0.0009590852,0.001158057,0.000593466,0.001287934,0.004110915],"category_scores_gemma":[0.022915566,0.00058200094,0.00047759383,0.00033439748,0.001072412,0.0015197449,0.0010890445,0.0010689431,0.0005481518],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012109541,0.058281656,0.2740655,0.0038065063,0.00042016312,0.003493128,0.35849103,0.0011622232,0.14244403,0.00089534116,0.0042087175,0.14062208],"study_design_scores_gemma":[0.00457684,0.10892689,0.7560031,0.00070835376,0.00066336326,0.0017249888,0.077995166,0.0062335283,0.029592143,0.0019365887,0.011137392,0.0005016855],"about_ca_topic_score_codex":0.0016048996,"about_ca_topic_score_gemma":0.0019142721,"teacher_disagreement_score":0.00491933,"about_ca_system_score_codex":0.00044288544,"about_ca_system_score_gemma":0.00073269533,"threshold_uncertainty_score":0.026016235},"labels":[],"label_agreement":null},{"id":"W2113986544","doi":"10.1109/wcre.2004.6","title":"A framework for the comparison of nested software decompositions","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Plug-in; Cluster analysis; Software; Nested set model; Data mining; Nested loop join; Reverse engineering; Theoretical computer science; Algorithm; Programming language; Artificial intelligence; Relational database","score_opus":0.04343940446552314,"score_gpt":0.3614693224927655,"score_spread":0.31802991802724234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113986544","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025667907,0.00012858621,0.9945741,0.000046580568,0.00003736021,0.00015056308,0.00019228482,0.0011572335,0.0011464349],"genre_scores_gemma":[0.034319654,0.000093327915,0.96400374,0.000034613004,0.00002780458,0.0004256461,0.0004711666,0.00027743305,0.00034665808],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97603863,0.0103658065,0.0027794659,0.0025052677,0.007426804,0.0008841026],"domain_scores_gemma":[0.9574206,0.023018403,0.0037699572,0.0070494707,0.007788696,0.00095293054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02393658,0.0021221128,0.0019290982,0.0108257895,0.0015225853,0.005146179,0.0030531643,0.0017985087,0.004829016],"category_scores_gemma":[0.064630136,0.0011066603,0.0023903684,0.0055556362,0.0030613816,0.0061298683,0.004933821,0.00290277,0.0011219793],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007947528,0.0004340304,0.0054159383,0.0010023686,0.00039966835,0.0004450687,0.0011142343,0.0660791,0.01960783,0.49507153,0.0072273063,0.40240824],"study_design_scores_gemma":[0.00025751954,0.0010870885,0.004643637,0.0005221432,0.00025557962,0.0008302359,0.00061029126,0.34912315,0.022030858,0.5608368,0.059543014,0.00025967002],"about_ca_topic_score_codex":0.002417114,"about_ca_topic_score_gemma":0.0021118482,"teacher_disagreement_score":0.02393658,"about_ca_system_score_codex":0.0014042566,"about_ca_system_score_gemma":0.0025480203,"threshold_uncertainty_score":0.12659025},"labels":[],"label_agreement":null},{"id":"W2114607209","doi":"10.1145/1107656.1107661","title":"Humans in the traceability loop","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University; National Aeronautics and Space Administration","keywords":"Traceability; TRACE (psycholinguistics); Computer science; Process (computing); Requirements traceability; Work (physics); Risk analysis (engineering); Software engineering; Process management; Engineering; Programming language; Business; Software; Requirements analysis","score_opus":0.02571180076463451,"score_gpt":0.2898022204966081,"score_spread":0.26409041973197356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114607209","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05191848,0.0020563025,0.7719911,0.064545445,0.0011687421,0.0006383885,0.00015080263,0.005011254,0.102519445],"genre_scores_gemma":[0.68140954,0.0020298108,0.26991996,0.013225778,0.00080330815,0.0006130533,0.00041984642,0.0011314974,0.030447254],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.953564,0.028325075,0.001259103,0.0052393195,0.009618826,0.001993723],"domain_scores_gemma":[0.8931534,0.06948625,0.005149056,0.015239923,0.011617778,0.0053534773],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03442513,0.0013823201,0.0008402478,0.0037288088,0.0045188633,0.0152319465,0.0034153117,0.0045499178,0.014100212],"category_scores_gemma":[0.0919747,0.0011672114,0.0006828054,0.0017384447,0.01590345,0.019305766,0.009610982,0.006964076,0.0051470064],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008716413,0.00075997604,0.016925521,0.0009385967,0.00032407956,0.0012790414,0.059096973,0.011791324,0.015980085,0.39206392,0.042655315,0.4573135],"study_design_scores_gemma":[0.00014670241,0.00038162753,0.002921127,0.00060736306,0.000102810234,0.000464937,0.009940018,0.018453673,0.0075893416,0.67077065,0.28841823,0.000203514],"about_ca_topic_score_codex":0.007018582,"about_ca_topic_score_gemma":0.002785989,"teacher_disagreement_score":0.03442513,"about_ca_system_score_codex":0.0027227974,"about_ca_system_score_gemma":0.009129333,"threshold_uncertainty_score":0.18205965},"labels":[],"label_agreement":null},{"id":"W2114622266","doi":"10.1109/icsmc.1994.400190","title":"Quantitative evaluation of expert systems","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Expert system; Quality (philosophy); Quality assurance; Data mining; Artificial intelligence; Engineering","score_opus":0.1592782373554071,"score_gpt":0.36450558182964926,"score_spread":0.20522734447424215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114622266","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33887997,0.0049764183,0.5827703,0.0015944849,0.0005995004,0.0020074754,0.0042057224,0.0018489559,0.06311715],"genre_scores_gemma":[0.8598167,0.0009246895,0.13088027,0.00017796725,0.00015934315,0.0012739097,0.0025566306,0.0001975364,0.004012936],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94131315,0.027039912,0.003962227,0.0022893318,0.024672894,0.0007224903],"domain_scores_gemma":[0.7787442,0.14775874,0.01378293,0.009796652,0.047469847,0.0024475649],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027271302,0.0010176551,0.0007737395,0.007088732,0.0006331113,0.003103412,0.0009777974,0.0010384545,0.0048695332],"category_scores_gemma":[0.16525082,0.00023653486,0.0004762541,0.0039200094,0.0013186315,0.0034435096,0.0014859147,0.0007042818,0.0007226532],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015630478,0.0008676211,0.05475971,0.0042384053,0.0007857242,0.0002222702,0.0035414943,0.11405998,0.021843279,0.08146509,0.01465757,0.70199585],"study_design_scores_gemma":[0.00042042663,0.0065442575,0.15562773,0.0016612045,0.00063030625,0.0009024958,0.004790247,0.55594563,0.053119354,0.12854885,0.091249526,0.0005600086],"about_ca_topic_score_codex":0.0010770166,"about_ca_topic_score_gemma":0.00102465,"teacher_disagreement_score":0.027271302,"about_ca_system_score_codex":0.0017580354,"about_ca_system_score_gemma":0.0014522313,"threshold_uncertainty_score":0.14422613},"labels":[],"label_agreement":null},{"id":"W2114760352","doi":"10.1109/isie.2006.295995","title":"Modeling Functional Requirements to Support Traceability Analysis","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Traceability; Requirements traceability; Computer science; Software engineering; Requirements analysis; Software requirements; Functional requirement; Software construction; Context (archaeology); Non-functional requirement; Software requirements specification; Verification and validation; Software system; Requirement; Software; Systems engineering; Programming language; Engineering","score_opus":0.04914235196886094,"score_gpt":0.29839395479459585,"score_spread":0.2492516028257349,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114760352","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004362169,0.00007520204,0.98729986,0.00022877796,0.000023258102,0.00020334822,0.00021569419,0.00082578126,0.006765867],"genre_scores_gemma":[0.15350884,0.00044561006,0.8375594,0.00014072005,0.000031287633,0.0008748713,0.0011386705,0.00037593706,0.005924697],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9963894,0.0013974971,0.00026563232,0.0003562479,0.0013776777,0.00021354032],"domain_scores_gemma":[0.99322563,0.0038791324,0.00068407186,0.0010689077,0.0010630226,0.0000791851],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030679666,0.0015101165,0.00053494493,0.0029497142,0.0005036374,0.0019508693,0.0021941117,0.0020669892,0.0061072595],"category_scores_gemma":[0.012878542,0.0008386037,0.0017399788,0.0016902754,0.0010418013,0.003212364,0.0014815158,0.0018429102,0.0012864402],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008337899,0.0001855563,0.001237425,0.0005553918,0.00008255028,0.0010197842,0.0016855849,0.3815983,0.012848882,0.5316831,0.0038446602,0.06517537],"study_design_scores_gemma":[0.00007065244,0.00011667796,0.00041298644,0.00021982317,0.00007371177,0.0003718731,0.00025177313,0.7600017,0.010743017,0.16428757,0.063388035,0.00006211304],"about_ca_topic_score_codex":0.007904808,"about_ca_topic_score_gemma":0.008214085,"teacher_disagreement_score":0.007904808,"about_ca_system_score_codex":0.0013197906,"about_ca_system_score_gemma":0.0019131973,"threshold_uncertainty_score":0.020430863},"labels":[],"label_agreement":null},{"id":"W2114826097","doi":"10.1109/icsm.2009.5306332","title":"Decomposing object-oriented class modules using an agglomerative clustering technique","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Cohesion (chemistry); Computer science; Jaccard index; Cluster analysis; Software quality; Class (philosophy); Data mining; Hierarchical clustering; Software engineering; Software; Software development; Artificial intelligence; Programming language","score_opus":0.025828260355329667,"score_gpt":0.3087353495919927,"score_spread":0.2829070892366631,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114826097","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028943827,0.00012788188,0.9673188,0.00006776623,0.000027478256,0.00048328374,0.00012711271,0.0016629493,0.0012409635],"genre_scores_gemma":[0.06252778,0.00010390035,0.93400574,0.00003057255,0.000017365495,0.00033508014,0.0006244148,0.00040676378,0.0019485172],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9978962,0.0003174929,0.00013633796,0.0003611716,0.0010846339,0.00020410711],"domain_scores_gemma":[0.9964573,0.000771999,0.00033529042,0.000611945,0.0016741992,0.00014920892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018924096,0.0015191524,0.0012821008,0.010326134,0.0021220127,0.0016793924,0.0019768078,0.00110321,0.0021241575],"category_scores_gemma":[0.0055961013,0.00080820225,0.0021253868,0.007343647,0.0008991603,0.0011960042,0.001384117,0.0008359031,0.0017150856],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003205013,0.0004914421,0.007230985,0.0005352885,0.0003875622,0.00036524882,0.0032117583,0.07135288,0.04211943,0.010764082,0.007030137,0.8561905],"study_design_scores_gemma":[0.00008996608,0.00026577458,0.015454204,0.00011585911,0.00038790124,0.0005928912,0.0017242966,0.88114417,0.03845776,0.030061912,0.031461388,0.0002439231],"about_ca_topic_score_codex":0.019677822,"about_ca_topic_score_gemma":0.020505395,"teacher_disagreement_score":0.019677822,"about_ca_system_score_codex":0.001194932,"about_ca_system_score_gemma":0.002520176,"threshold_uncertainty_score":0.039126575},"labels":[],"label_agreement":null},{"id":"W2114836483","doi":"10.1002/smr.421","title":"Identification of behavioural and creational design motifs through dynamic analysis","year":2009,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Identification (biology); Visitor pattern; Computer science; Java; Program comprehension; Software engineering; Constraint (computer-aided design); Software; Theoretical computer science; Programming language; Engineering; Software system","score_opus":0.0626294758113364,"score_gpt":0.36130520564859697,"score_spread":0.29867572983726054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114836483","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06784951,0.00013734466,0.92539567,0.00016282962,0.000014866333,0.00016210155,0.00026148706,0.0028593051,0.003156956],"genre_scores_gemma":[0.35052803,0.00015168663,0.64629096,0.00004945473,0.000010507672,0.0002504717,0.000776145,0.0003767023,0.0015661002],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99745446,0.00065697957,0.00016304756,0.00053358014,0.0010106286,0.00018139344],"domain_scores_gemma":[0.99066556,0.004394058,0.0015992485,0.0015573163,0.0016114397,0.0001724046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018998473,0.0008190169,0.00057085,0.004286665,0.0006930415,0.0018546432,0.0013782168,0.00094792957,0.0022328172],"category_scores_gemma":[0.0119846035,0.00072632247,0.0008595282,0.0018829411,0.0012312704,0.0018123516,0.001515992,0.0010552955,0.00045939191],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000374624,0.0003387578,0.03273567,0.0008648635,0.00017275904,0.000938251,0.0019204458,0.16494267,0.049278934,0.13227707,0.004375567,0.6117804],"study_design_scores_gemma":[0.000027760123,0.000054442153,0.0028042574,0.00008328478,0.000051881438,0.0004227602,0.00022946724,0.9322368,0.01985848,0.033496026,0.010688258,0.00004647988],"about_ca_topic_score_codex":0.0039012325,"about_ca_topic_score_gemma":0.0043763113,"teacher_disagreement_score":0.004286665,"about_ca_system_score_codex":0.0012559909,"about_ca_system_score_gemma":0.0017249091,"threshold_uncertainty_score":0.010047495},"labels":[],"label_agreement":null},{"id":"W2114839529","doi":"10.1109/wcre.2009.25","title":"A Study of the Time Dependence of Code Changes","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Documentation; Computer science; Period (music); Code (set theory); Open source; Source code; Foundation (evidence); Data science; Software; Software engineering; History; Programming language; Archaeology","score_opus":0.027138251579830658,"score_gpt":0.28387888478050055,"score_spread":0.25674063320066987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114839529","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98131114,0.0005536112,0.012719065,0.00024034981,0.000011100992,0.00004366176,0.00052400085,0.00012361897,0.004473361],"genre_scores_gemma":[0.9955597,0.00019153327,0.0031817385,0.00002365464,0.000015621943,0.000033510325,0.00044426956,0.000046195055,0.0005036472],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9978752,0.00066159654,0.00016014853,0.0005119447,0.00063888705,0.00015222626],"domain_scores_gemma":[0.9129406,0.058858458,0.017321395,0.0045805606,0.004805113,0.0014937536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026584552,0.00023811407,0.00031502036,0.0038379673,0.0004945662,0.0009808985,0.0006285618,0.0005468232,0.0015890441],"category_scores_gemma":[0.059303418,0.000398148,0.00039177126,0.004502593,0.0008099986,0.0024899004,0.00073869317,0.0010186441,0.00030211874],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006164451,0.00023680126,0.8552126,0.00018562477,0.0002202865,0.0004685367,0.00425496,0.016467106,0.009573369,0.015407263,0.0011574716,0.09619957],"study_design_scores_gemma":[0.000020952284,0.00020419307,0.9351299,0.000039822495,0.0000674163,0.00045189634,0.000832853,0.050025426,0.0022722862,0.0069340775,0.0039685536,0.0000526427],"about_ca_topic_score_codex":0.007904426,"about_ca_topic_score_gemma":0.0052234973,"teacher_disagreement_score":0.007904426,"about_ca_system_score_codex":0.0009076401,"about_ca_system_score_gemma":0.0005262423,"threshold_uncertainty_score":0.01571685},"labels":[],"label_agreement":null},{"id":"W2114905813","doi":"10.5555/2337223.2337410","title":"CodeTimeline: storytelling with versioning data","year":2012,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Storytelling; Software versioning; Casual; World Wide Web; Visualization; Data visualization; Narrative; Human–computer interaction; Software engineering; Software; Programming language; Artificial intelligence","score_opus":0.06457803039089124,"score_gpt":0.2980390296624419,"score_spread":0.23346099927155065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114905813","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04582313,0.0005240447,0.86667144,0.0010035904,0.00021092684,0.0010043376,0.0049140844,0.06427798,0.015570461],"genre_scores_gemma":[0.27726948,0.0005217032,0.6907271,0.0003361814,0.00012608481,0.001740415,0.008002547,0.008298781,0.0129776],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99827445,0.0009872529,0.00010410458,0.00024119044,0.00031463648,0.00007830331],"domain_scores_gemma":[0.98436326,0.012141973,0.00054809306,0.001867486,0.0005553344,0.0005238785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00268626,0.0015629043,0.00048935984,0.0015225643,0.0007646963,0.0026432727,0.0022055504,0.0015059621,0.019870032],"category_scores_gemma":[0.017463204,0.0007223636,0.0006662653,0.00089425815,0.0009198472,0.0049456973,0.0038389515,0.0015081054,0.0026078136],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018051075,0.0006925716,0.0063318815,0.0034316082,0.00022289158,0.0021318356,0.04634281,0.015637895,0.050616663,0.028729476,0.09978139,0.74427587],"study_design_scores_gemma":[0.00074394513,0.001429227,0.006947925,0.0008662237,0.000232429,0.0025373052,0.0077790283,0.20440328,0.06761473,0.054356493,0.65267545,0.00041401372],"about_ca_topic_score_codex":0.00091456313,"about_ca_topic_score_gemma":0.0018492269,"teacher_disagreement_score":0.019870032,"about_ca_system_score_codex":0.0003888562,"about_ca_system_score_gemma":0.00037596884,"threshold_uncertainty_score":0.066471934},"labels":[],"label_agreement":null},{"id":"W2114928940","doi":"10.1016/j.datak.2013.02.001","title":"Assessing the quality factors found in in-line documentation written in natural language: The JavadocMiner","year":2013,"lang":"en","type":"article","venue":"Data & Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"U.S. Department of Defense","keywords":"Documentation; Natural (archaeology); Quality (philosophy); Line (geometry); Natural language processing; Computer science; Linguistics; Psychology; History; Mathematics; Programming language; Archaeology; Philosophy","score_opus":0.059091997112096124,"score_gpt":0.39128965147035494,"score_spread":0.3321976543582588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114928940","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98226905,0.00027697804,0.012053192,0.00014868383,0.000016325901,0.00006660742,0.0011350025,0.0028184643,0.0012156296],"genre_scores_gemma":[0.91683245,0.00022007029,0.074730635,0.00007109977,0.000014660435,0.00007035029,0.0053570233,0.0010028703,0.0017008967],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9932473,0.0017380164,0.0012451132,0.00086395716,0.002699762,0.0002058676],"domain_scores_gemma":[0.8842165,0.0609452,0.018629316,0.010188616,0.024601066,0.0014192596],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008280525,0.00045649408,0.0005453289,0.0036784825,0.0004767217,0.0021369196,0.00080319954,0.00086806825,0.0009938509],"category_scores_gemma":[0.06522698,0.0003589376,0.00037473039,0.0026500318,0.00036615558,0.0016488519,0.0011349887,0.0006043764,0.0005257657],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023562724,0.001920891,0.3902852,0.0025673616,0.0004928734,0.00082413305,0.0065176887,0.0063798446,0.073204316,0.0010415744,0.008980397,0.50542957],"study_design_scores_gemma":[0.00062999857,0.0038094698,0.6573953,0.00084212964,0.0010794943,0.0026403018,0.0052062497,0.12965006,0.1717581,0.0019633737,0.024640746,0.00038470884],"about_ca_topic_score_codex":0.0028108067,"about_ca_topic_score_gemma":0.0049648476,"teacher_disagreement_score":0.008280525,"about_ca_system_score_codex":0.00059099926,"about_ca_system_score_gemma":0.0015801025,"threshold_uncertainty_score":0.04379213},"labels":[],"label_agreement":null},{"id":"W2114938579","doi":"10.1145/1512475.1512493","title":"Towards a more efficient static software change impact analysis method","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Change impact analysis; Computer science; Static analysis; Software; Software testing; Software engineering; Programming language","score_opus":0.05565394329737228,"score_gpt":0.3614763793535488,"score_spread":0.30582243605617654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114938579","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008526945,0.00009875927,0.9853088,0.000118473836,0.0000248837,0.00011567776,0.00020996135,0.004803923,0.00079252146],"genre_scores_gemma":[0.11772418,0.00014253652,0.8793859,0.0000687855,0.000051842955,0.0002424762,0.0007829277,0.00034453318,0.0012568851],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964874,0.0005001023,0.00020597602,0.0005153557,0.0021542,0.0001369949],"domain_scores_gemma":[0.99332434,0.0019845,0.0005754676,0.00103093,0.002976127,0.00010863414],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019322801,0.0012391938,0.0013942616,0.0076983417,0.0007459067,0.0018909177,0.001822305,0.0009738531,0.0016983569],"category_scores_gemma":[0.010681314,0.0006508761,0.0011666068,0.0038300431,0.0005801053,0.002067242,0.0012299207,0.0012871957,0.0010278862],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013311557,0.00023593554,0.008828396,0.00023130066,0.00019958577,0.00019205843,0.00030642565,0.07125079,0.017881203,0.0104187345,0.005348048,0.88497454],"study_design_scores_gemma":[0.000036779973,0.00008463674,0.0045532393,0.0000433312,0.00008020004,0.00027309952,0.00017527293,0.9549422,0.018014777,0.014529016,0.007207095,0.00006041665],"about_ca_topic_score_codex":0.0048440923,"about_ca_topic_score_gemma":0.004227918,"teacher_disagreement_score":0.0076983417,"about_ca_system_score_codex":0.00107645,"about_ca_system_score_gemma":0.002023113,"threshold_uncertainty_score":0.010218978},"labels":[],"label_agreement":null},{"id":"W2115057462","doi":"10.1109/wcre.2004.16","title":"Combined software and hardware comprehension in reverse engineering","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Porting; Computer science; Interpreter; Software; Software engineering; Process (computing); Backporting; Reverse engineering; Program comprehension; Software maintenance; Comprehension; Operating system; Software development; Embedded system; Software development process; Software system; Programming language","score_opus":0.01196586451773886,"score_gpt":0.2291831265115527,"score_spread":0.21721726199381383,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115057462","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41448638,0.0012684827,0.5239748,0.0042787846,0.00009399039,0.00024008317,0.00002174487,0.0015034642,0.054132316],"genre_scores_gemma":[0.85717773,0.0007289665,0.12687282,0.000741412,0.00005082853,0.000120677534,0.00005693905,0.00048159045,0.013769006],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9898331,0.0073436494,0.00020039674,0.000592932,0.0015183206,0.0005116207],"domain_scores_gemma":[0.9740757,0.018960152,0.001123107,0.003333165,0.0021345941,0.00037328523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007061585,0.0011361251,0.00064617896,0.0010860307,0.0013443936,0.0032480222,0.0013720302,0.002116215,0.004428952],"category_scores_gemma":[0.029448245,0.000646138,0.0004613138,0.0005782697,0.0054148934,0.008967858,0.0046952916,0.0026771163,0.0010831243],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045689472,0.001065752,0.0073991613,0.00090903806,0.000056337376,0.0043269666,0.2471608,0.0075250263,0.056349885,0.0988558,0.005249724,0.5706446],"study_design_scores_gemma":[0.0003389285,0.0042205867,0.010786788,0.0013905102,0.00024822922,0.023164129,0.15510876,0.074707106,0.13417646,0.34092405,0.25441757,0.0005168666],"about_ca_topic_score_codex":0.0007362429,"about_ca_topic_score_gemma":0.0011824521,"teacher_disagreement_score":0.007061585,"about_ca_system_score_codex":0.0007358866,"about_ca_system_score_gemma":0.0017707389,"threshold_uncertainty_score":0.037345707},"labels":[],"label_agreement":null},{"id":"W2115081444","doi":"10.1109/scam.2008.9","title":"Using Program Transformations to Add Structure to a Legacy Data Model","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Programming language; Programmer; Java; Legacy system; Data modeling; Data model (GIS); Set (abstract data type); Relocation; Overlay; Database; Software; Artificial intelligence","score_opus":0.15012639048307888,"score_gpt":0.3698353484788266,"score_spread":0.21970895799574772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115081444","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020420577,0.00007261092,0.96024114,0.00043065127,0.00011166228,0.00019079886,0.0002914117,0.011894485,0.0063466574],"genre_scores_gemma":[0.16048129,0.00036710742,0.8247474,0.00041389937,0.00006567912,0.0002650696,0.0013332368,0.0067324503,0.0055938293],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99849427,0.0002933503,0.00012252964,0.00025148265,0.0007290471,0.00010933335],"domain_scores_gemma":[0.9931265,0.0017247698,0.00043815444,0.0037779147,0.0008336796,0.00009892875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020957123,0.0006928863,0.00035314492,0.0009864838,0.0006231142,0.0020183788,0.001388407,0.0006734527,0.0022894158],"category_scores_gemma":[0.011288457,0.00066308945,0.0009699168,0.00077946414,0.0013870014,0.003242548,0.0028095574,0.0022735319,0.001512687],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003719738,0.0006718016,0.013160328,0.0009982385,0.00029731222,0.002092778,0.0048812567,0.04996903,0.0881917,0.31652236,0.024740102,0.49810314],"study_design_scores_gemma":[0.0000959355,0.00025618225,0.0028840601,0.0003919833,0.00023727636,0.0017193333,0.00075588794,0.18251522,0.20193878,0.22343576,0.38558084,0.0001887388],"about_ca_topic_score_codex":0.001250076,"about_ca_topic_score_gemma":0.0018289342,"teacher_disagreement_score":0.0022894158,"about_ca_system_score_codex":0.0006644736,"about_ca_system_score_gemma":0.0019261809,"threshold_uncertainty_score":0.011083305},"labels":[],"label_agreement":null},{"id":"W2115130131","doi":"10.1145/2568225.2568313","title":"Live API documentation","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":248,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Documentation; Computer science; Application programming interface; Abstraction; World Wide Web; Software engineering; Programming language","score_opus":0.01053191051377647,"score_gpt":0.2687611003884102,"score_spread":0.2582291898746338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115130131","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004982883,0.0016376713,0.4765014,0.004487249,0.0023331225,0.00074379705,0.026275031,0.17474028,0.3082986],"genre_scores_gemma":[0.070811324,0.0041819005,0.32146883,0.0035494557,0.0012559994,0.0014838183,0.09079582,0.10300879,0.40344408],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99400795,0.00111859,0.00068411196,0.0005013617,0.0032502834,0.00043769222],"domain_scores_gemma":[0.96787274,0.007589638,0.0011230693,0.0119552445,0.010508572,0.00095079426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005371414,0.0013247732,0.00078806496,0.0041419654,0.0015483473,0.0061073923,0.0035317321,0.0028465632,0.15813197],"category_scores_gemma":[0.03517716,0.0012305417,0.0008327878,0.003240456,0.0008722121,0.009530368,0.0062072487,0.005008946,0.11675573],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019139166,0.00014936077,0.00067295186,0.00083006977,0.000023308696,0.0003980066,0.0010633331,0.00096935197,0.0037669183,0.04702592,0.6006026,0.34430674],"study_design_scores_gemma":[0.00003096116,0.000023628274,0.00030490296,0.00031341077,0.000010988085,0.00046151865,0.00011106286,0.0018170248,0.0042617545,0.009402862,0.98322624,0.000035686167],"about_ca_topic_score_codex":0.0022229552,"about_ca_topic_score_gemma":0.0029517442,"teacher_disagreement_score":0.15813197,"about_ca_system_score_codex":0.0011833524,"about_ca_system_score_gemma":0.0032842958,"threshold_uncertainty_score":0.52900416},"labels":[],"label_agreement":null},{"id":"W2115185340","doi":"10.1109/icsc.2009.19","title":"Using a Formal Language Constructs for Software Model Evolution","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; Carleton University","funders":"European Research Consortium for Informatics and Mathematics","keywords":"Computer science; Unified Modeling Language; Object Constraint Language; Programming language; Class diagram; Dependency (UML); Directed acyclic graph; Graph; Dependency graph; Theoretical computer science; Software; Applications of UML; Software engineering; Algorithm","score_opus":0.03217343751371565,"score_gpt":0.30476991119752994,"score_spread":0.2725964736838143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115185340","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0007827562,0.00004262449,0.99682957,0.00023335866,0.000033866963,0.000047660473,0.00004187355,0.0011240838,0.0008642663],"genre_scores_gemma":[0.025922684,0.00018036757,0.9718075,0.00017177081,0.000040177874,0.00018539406,0.00019443831,0.0004518074,0.0010459028],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964719,0.0011786906,0.0004688009,0.00045946892,0.0012415628,0.00017956532],"domain_scores_gemma":[0.99154085,0.004475684,0.0007032783,0.0020496792,0.00106302,0.00016755091],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004665446,0.0010429382,0.00045629524,0.0013813127,0.00074036245,0.002397236,0.001825439,0.0011410938,0.0029219047],"category_scores_gemma":[0.01120796,0.0007902119,0.0016816612,0.0008466613,0.0029083232,0.004950226,0.002276697,0.0034018382,0.00105388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005280566,0.00013643374,0.0006178841,0.0006150721,0.0000510379,0.0004020867,0.0011977985,0.01954502,0.019166538,0.864003,0.004177993,0.09003424],"study_design_scores_gemma":[0.00017859219,0.00017413744,0.00034015748,0.00081851985,0.0001389382,0.0012004448,0.0002626893,0.21519037,0.053224742,0.36707562,0.36121112,0.00018469551],"about_ca_topic_score_codex":0.00194071,"about_ca_topic_score_gemma":0.0015518826,"teacher_disagreement_score":0.004665446,"about_ca_system_score_codex":0.0013048459,"about_ca_system_score_gemma":0.0026698401,"threshold_uncertainty_score":0.024673522},"labels":[],"label_agreement":null},{"id":"W2115398138","doi":"10.1109/wpc.2002.1021334","title":"An integrated approach for studying architectural evolution","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":97,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software evolution; Architectural pattern; Visualization; Software engineering; Software; Software architecture; Extensibility; Software system; Software visualization; Component-based software engineering; Software construction; Data mining; Programming language","score_opus":0.03115702757862012,"score_gpt":0.2787823887526455,"score_spread":0.24762536117402537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115398138","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010071306,0.0002816198,0.9843162,0.00026571137,0.000018120205,0.00007252998,0.00012869466,0.0012035653,0.003642282],"genre_scores_gemma":[0.069050334,0.00047715282,0.9282995,0.000044891633,0.000019949648,0.0001885725,0.00030739745,0.00014237405,0.0014698093],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9988404,0.000319842,0.00009855571,0.00021986192,0.00047862646,0.000042645595],"domain_scores_gemma":[0.9980983,0.00060925534,0.00026218256,0.000545617,0.00037913406,0.00010553545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014797801,0.0008496997,0.0007421761,0.0045474456,0.0005791575,0.0022665262,0.0014580076,0.0010142737,0.0025645543],"category_scores_gemma":[0.004493676,0.0004993315,0.00084259256,0.0041559734,0.0010389036,0.006033788,0.0020010704,0.0015124097,0.00042352386],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009165254,0.00036319075,0.015785761,0.00059264817,0.00031897207,0.0006098523,0.004031008,0.057977047,0.046281517,0.43645227,0.004402413,0.43309367],"study_design_scores_gemma":[0.0000686933,0.0005192871,0.016128771,0.00027512613,0.0003516934,0.0018960617,0.0016682717,0.3823863,0.017211232,0.46506247,0.114255585,0.00017648049],"about_ca_topic_score_codex":0.0023474903,"about_ca_topic_score_gemma":0.0029426685,"teacher_disagreement_score":0.0045474456,"about_ca_system_score_codex":0.0007944052,"about_ca_system_score_gemma":0.0011247322,"threshold_uncertainty_score":0.008579254},"labels":[],"label_agreement":null},{"id":"W2115486982","doi":"10.1109/csmr.2001.914988","title":"Strategies for migration from C to Java","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Java; Real time Java; Programming language; Java annotation; Generics in Java; Popularity; strictfp","score_opus":0.026287166467984655,"score_gpt":0.28456310131588936,"score_spread":0.2582759348479047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115486982","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03846984,0.0014553188,0.88903856,0.004008158,0.0010973904,0.0010847243,0.00014718572,0.014238283,0.050460454],"genre_scores_gemma":[0.16216877,0.0015791041,0.81064117,0.0024350767,0.00029045306,0.00085336424,0.00046282413,0.0044205547,0.017148783],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965469,0.0008005512,0.0003264626,0.00048572256,0.0012971336,0.00054319727],"domain_scores_gemma":[0.99117893,0.0018004613,0.000411934,0.0033769067,0.0028418417,0.0003899501],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028381932,0.001088332,0.00056758424,0.0020386386,0.0014296841,0.0030434409,0.0026263716,0.0014328005,0.004482529],"category_scores_gemma":[0.01834492,0.0008237898,0.0009483407,0.0017568233,0.0009237794,0.0036021976,0.003343471,0.0027543015,0.003532875],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058227876,0.00044237063,0.0037747014,0.0005776675,0.00010276762,0.0013143077,0.0034604073,0.011531093,0.04000503,0.11744787,0.044146832,0.77661467],"study_design_scores_gemma":[0.0004092579,0.00073719444,0.004007625,0.00087956694,0.00025599592,0.002655148,0.0025652442,0.13448845,0.07791079,0.15078892,0.6249518,0.0003500511],"about_ca_topic_score_codex":0.0034407286,"about_ca_topic_score_gemma":0.0024807318,"teacher_disagreement_score":0.004482529,"about_ca_system_score_codex":0.0007305594,"about_ca_system_score_gemma":0.0025910644,"threshold_uncertainty_score":0.015009999},"labels":[],"label_agreement":null},{"id":"W2115513993","doi":"10.1109/wcre.2011.26","title":"Useful, But Usable? Factors Affecting the Usability of APIs","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Usability; Computer science; Pluralistic walkthrough; USable; Cognitive walkthrough; Usability inspection; Usability engineering; Usability lab; Usability goals; World Wide Web; Web usability; Quality (philosophy); Heuristic evaluation; Software engineering; Human–computer interaction","score_opus":0.06521488158642712,"score_gpt":0.2625446144278703,"score_spread":0.19732973284144317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115513993","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9939751,0.00038927197,0.0014779207,0.00033490185,0.000013827026,0.000048210095,0.0000337647,0.00003545161,0.003691505],"genre_scores_gemma":[0.99834883,0.00017113832,0.0010031512,0.00005279272,0.000013876669,0.000024055475,0.000031922285,0.000026049342,0.00032808096],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.98034817,0.0091662705,0.002148597,0.0009956086,0.006119966,0.0012213708],"domain_scores_gemma":[0.6508703,0.27296868,0.040397074,0.009395527,0.021764869,0.004603611],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015643304,0.00058640534,0.0005051714,0.0027740572,0.0010526377,0.0054879813,0.00066590693,0.0010705446,0.0014956369],"category_scores_gemma":[0.19313702,0.0005103107,0.00055277755,0.0021880392,0.002365036,0.0033193137,0.0013144097,0.001217023,0.00037864203],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034742476,0.00041162758,0.8864298,0.0005020605,0.00018913185,0.0005514567,0.027581988,0.0003441771,0.0037821173,0.0009637793,0.00095787365,0.07793857],"study_design_scores_gemma":[0.00001987068,0.00039988133,0.9703358,0.0001916852,0.0001460828,0.0006454099,0.023057815,0.0012178201,0.0008247341,0.00078531663,0.002319234,0.000056305653],"about_ca_topic_score_codex":0.0031041866,"about_ca_topic_score_gemma":0.0034841741,"teacher_disagreement_score":0.015643304,"about_ca_system_score_codex":0.0008355977,"about_ca_system_score_gemma":0.0013911831,"threshold_uncertainty_score":0.08273065},"labels":[],"label_agreement":null},{"id":"W2115571575","doi":"10.1109/scam.2001.972678","title":"Software engineering by source transformation - experience with TXL","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; University of Waterloo; Queen's University","funders":"","keywords":"Computer science; Programming language; Software engineering; COBOL; Software construction; Software maintenance; Software development; Program transformation; Source code; Code refactoring; Backporting; Software; Operating system","score_opus":0.01048646824423637,"score_gpt":0.20291743342019095,"score_spread":0.19243096517595457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115571575","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015094599,0.0066055553,0.872745,0.014198901,0.0006473141,0.000104519546,0.000110406116,0.008346503,0.08214736],"genre_scores_gemma":[0.17848161,0.019286362,0.72144765,0.006363304,0.001036719,0.00022203718,0.0007536847,0.007607995,0.06480075],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9923707,0.003152185,0.00048162288,0.0008590336,0.0028742796,0.00026212836],"domain_scores_gemma":[0.97756505,0.012255191,0.0005719312,0.005765791,0.0029415067,0.0009005638],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010624019,0.00064406544,0.0005342919,0.0013584121,0.0011863032,0.006074069,0.0019226017,0.0017799939,0.00677223],"category_scores_gemma":[0.027628023,0.00069521356,0.0007687327,0.0026613076,0.0040732143,0.0153037235,0.0055310666,0.006827996,0.0053683654],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010454301,0.0002710608,0.0014535915,0.00042468644,0.000024622706,0.0005693997,0.019871054,0.003948064,0.0074070906,0.22323626,0.04973911,0.69295055],"study_design_scores_gemma":[0.000038259135,0.00013327478,0.00048026943,0.00029645977,0.000011877812,0.0017188344,0.001742733,0.009267736,0.008541601,0.1486102,0.8290885,0.00007025907],"about_ca_topic_score_codex":0.00082978845,"about_ca_topic_score_gemma":0.0007443208,"teacher_disagreement_score":0.010624019,"about_ca_system_score_codex":0.0015026653,"about_ca_system_score_gemma":0.0022031341,"threshold_uncertainty_score":0.05618584},"labels":[],"label_agreement":null},{"id":"W2115655439","doi":"10.1109/icis-comsar.2006.34","title":"Empirical Validation of Test-Driven Pair Programming in Game Development","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Algoma University; Laurentian University","funders":"","keywords":"Code refactoring; Pair programming; Extreme programming; Computer science; Cohesion (chemistry); Extreme programming practices; Waterfall model; Video game development; Test (biology); Test-driven development; Programming language; Software engineering; Game design; Software development; Human–computer interaction; Software development process; Software","score_opus":0.02931252996577616,"score_gpt":0.29357458445614226,"score_spread":0.2642620544903661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115655439","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9903812,0.000035965608,0.007599372,0.00006704584,0.0000124180315,0.00048857636,0.000033054675,0.000029126619,0.0013533355],"genre_scores_gemma":[0.98671174,0.000043223677,0.0117927855,0.000075185926,0.000009156785,0.0007664394,0.000114116,0.00002656364,0.00046080793],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.95322114,0.034085788,0.0020915018,0.0032945657,0.006283353,0.0010236565],"domain_scores_gemma":[0.5669676,0.34461468,0.018116856,0.032798875,0.030396175,0.007105803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03382371,0.0007192138,0.0005311655,0.0011992884,0.0010871104,0.0015673132,0.002595943,0.0013669508,0.0012968458],"category_scores_gemma":[0.22113839,0.00069338206,0.00052869326,0.00077127817,0.0022603665,0.0023550282,0.0030161238,0.0019368596,0.0004193258],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0068851924,0.049259864,0.5833107,0.0013548125,0.00064498564,0.0017582307,0.059682082,0.014953914,0.02365222,0.005266728,0.0015466257,0.25168458],"study_design_scores_gemma":[0.0024969822,0.09992028,0.6934092,0.0007575085,0.000473705,0.0031436374,0.033056207,0.07850996,0.06105922,0.007498109,0.019297594,0.00037759062],"about_ca_topic_score_codex":0.0011653872,"about_ca_topic_score_gemma":0.0013943359,"teacher_disagreement_score":0.03382371,"about_ca_system_score_codex":0.0012559096,"about_ca_system_score_gemma":0.0015207379,"threshold_uncertainty_score":0.17887902},"labels":[],"label_agreement":null},{"id":"W2115796408","doi":"10.1109/mise.2007.2","title":"A Framework for Empirical Evaluation of Model Comprehensibility","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Key (lock); Quality (philosophy); Empirical research; Grounded theory; Human–computer interaction; Qualitative research","score_opus":0.22764436621948622,"score_gpt":0.4608701870562838,"score_spread":0.23322582083679758,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115796408","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07151946,0.000835301,0.89212894,0.0021743204,0.0001126819,0.003768027,0.0008574775,0.00054052426,0.028063318],"genre_scores_gemma":[0.5533111,0.00030316296,0.42342427,0.00053075445,0.000053343603,0.020647688,0.00090037216,0.00024215365,0.00058710686],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.69487035,0.23759155,0.016565042,0.013614867,0.0344692,0.002888996],"domain_scores_gemma":[0.36086184,0.5231604,0.030462032,0.058593526,0.024947196,0.0019749682],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.23465468,0.003201296,0.0020609892,0.019492716,0.0034051263,0.011886062,0.0044197775,0.0047173975,0.0055993823],"category_scores_gemma":[0.59806263,0.0014221931,0.0034272203,0.014539507,0.017501898,0.02189794,0.01330298,0.0066319276,0.00065414247],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055760477,0.0015926876,0.047647074,0.0021058219,0.00082246057,0.00027898912,0.032006178,0.019778907,0.0037804777,0.7226422,0.0055688713,0.16321872],"study_design_scores_gemma":[0.00044493566,0.0016786205,0.039284922,0.0023481941,0.00030493503,0.0005241261,0.015462989,0.117545985,0.0035694658,0.7995271,0.018821336,0.0004873766],"about_ca_topic_score_codex":0.0029219587,"about_ca_topic_score_gemma":0.0018959237,"teacher_disagreement_score":0.23465468,"about_ca_system_score_codex":0.0060465764,"about_ca_system_score_gemma":0.0051611164,"threshold_uncertainty_score":0.943807},"labels":[],"label_agreement":null},{"id":"W2115993542","doi":"10.5381/jot.2004.3.4.a8","title":"A Proposal of a New Class Cohesion Criterion: An Empirical Study.","year":2004,"lang":"en","type":"article","venue":"The Journal of Object Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":81,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada; Uniwersytet Medyczny im. Karola Marcinkowskiego w Poznaniu","keywords":"Cohesion (chemistry); Class (philosophy); Computer science; Empirical research; Mathematics; Artificial intelligence; Statistics; Physics","score_opus":0.025927782708854038,"score_gpt":0.33584671452188414,"score_spread":0.3099189318130301,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115993542","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48892185,0.0023103708,0.4851016,0.0033215003,0.0003110078,0.0015622214,0.0013459777,0.0006418808,0.016483692],"genre_scores_gemma":[0.8574943,0.00047165915,0.1382661,0.00026453164,0.00011116097,0.0012409012,0.0011569891,0.0001492497,0.0008451403],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9728382,0.012343374,0.002231284,0.0035629405,0.008518849,0.00050532684],"domain_scores_gemma":[0.861934,0.09353057,0.009645831,0.006443217,0.025906568,0.0025398452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025337145,0.0010043125,0.0010244596,0.00877382,0.0013315412,0.0041560936,0.0022846262,0.0021128142,0.0025216928],"category_scores_gemma":[0.14758286,0.00041248364,0.0007824331,0.007964605,0.0027914762,0.011886146,0.0027453846,0.0016216655,0.0007078504],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003836673,0.0018474379,0.48096105,0.001773142,0.00043675856,0.00036102306,0.008949982,0.0069980207,0.00843671,0.05407069,0.014267588,0.42151392],"study_design_scores_gemma":[0.00040767653,0.0042328974,0.5967448,0.001247559,0.00076125946,0.0021022186,0.018320836,0.22050995,0.010081786,0.10044412,0.04459213,0.0005548203],"about_ca_topic_score_codex":0.0020895596,"about_ca_topic_score_gemma":0.0021131046,"teacher_disagreement_score":0.025337145,"about_ca_system_score_codex":0.0022018712,"about_ca_system_score_gemma":0.002230111,"threshold_uncertainty_score":0.1339972},"labels":[],"label_agreement":null},{"id":"W2116008481","doi":"10.1109/wcre.2008.48","title":"Error Correcting Graph Matching Application to Software Evolution","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Theoretical computer science; Graph; Software; Matching (statistics); Programming language; Mathematics","score_opus":0.018584682276340348,"score_gpt":0.26564369047189773,"score_spread":0.24705900819555737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116008481","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019257255,0.0001289388,0.97809565,0.00021704941,0.000049035607,0.000057213954,0.000022939563,0.0011581656,0.0010137107],"genre_scores_gemma":[0.3360895,0.00021391,0.6604517,0.00013471079,0.000038939266,0.000096930504,0.000113794114,0.00034794226,0.0025126005],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978003,0.00081181753,0.00012424313,0.0005217621,0.0006166424,0.00012513259],"domain_scores_gemma":[0.9912499,0.005216604,0.000670583,0.0015189758,0.0012040314,0.00013993213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021472068,0.0006827742,0.00094066863,0.002187528,0.00086741796,0.0009608434,0.001628431,0.0018452518,0.0017115689],"category_scores_gemma":[0.015479553,0.00030001393,0.0007288907,0.0028021846,0.0013788029,0.0018290264,0.001504589,0.001206453,0.00033169863],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001911094,0.00016845026,0.0026540847,0.00024676026,0.00010252006,0.00043133917,0.00042918167,0.42382085,0.022271987,0.092816874,0.0022402513,0.45462653],"study_design_scores_gemma":[0.000021004593,0.00004825275,0.00040461798,0.000014751555,0.000030630505,0.00022389859,0.000043295004,0.92043406,0.016908186,0.05829916,0.003552693,0.00001953682],"about_ca_topic_score_codex":0.0034193192,"about_ca_topic_score_gemma":0.0022615078,"teacher_disagreement_score":0.0034193192,"about_ca_system_score_codex":0.0012464961,"about_ca_system_score_gemma":0.0013820567,"threshold_uncertainty_score":0.011355698},"labels":[],"label_agreement":null},{"id":"W2116054015","doi":"10.1109/scam.2015.7335409","title":"The impact of cross-distribution bug duplicates, empirical study on Debian and Ubuntu","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Precision and recall; Software bug; Distribution (mathematics); Work (physics); Open source; Database; Software; Operating system; Information retrieval; Engineering","score_opus":0.06886057734100762,"score_gpt":0.4118321382624069,"score_spread":0.34297156092139924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116054015","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9959657,0.00060627126,0.0019691861,0.000113514616,0.000011887989,0.00004434529,0.0002700466,0.00021877761,0.0008001949],"genre_scores_gemma":[0.9951881,0.00019203304,0.0033302512,0.00004849627,0.000009013158,0.00003299531,0.0007490604,0.00007229755,0.00037769345],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98462534,0.00494791,0.001478789,0.002463026,0.005784064,0.00070084457],"domain_scores_gemma":[0.7806507,0.14332616,0.03202522,0.016948758,0.023974683,0.003074497],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012368596,0.00053402386,0.0005604455,0.0042600404,0.0011740979,0.0015372239,0.0013006272,0.0007242175,0.00065596006],"category_scores_gemma":[0.124926016,0.0005323295,0.0004916703,0.0037750134,0.0014621491,0.0032326241,0.0018130972,0.0010452111,0.0002603872],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008513984,0.0006697693,0.8754213,0.00075716566,0.00040511385,0.001126925,0.005601907,0.0121630225,0.0029361045,0.0012548664,0.0038773904,0.09493495],"study_design_scores_gemma":[0.00011410413,0.001386504,0.9059584,0.0003466553,0.00036869742,0.0045034713,0.006579281,0.05694664,0.012788828,0.0016781957,0.009198868,0.00013028628],"about_ca_topic_score_codex":0.009703247,"about_ca_topic_score_gemma":0.01038529,"teacher_disagreement_score":0.9876314,"about_ca_system_score_codex":0.0016051196,"about_ca_system_score_gemma":0.0013794518,"threshold_uncertainty_score":0.065412164},"labels":[],"label_agreement":null},{"id":"W2116146752","doi":"10.1109/icci.2004.28","title":"Release planning under fuzzy effort constraints","year":2004,"lang":"en","type":"article","venue":"IEEE International Conference on Cognitive Informatics","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Fuzzy logic; Computer science; Stakeholder; Quality (philosophy); Fuzzy set; Product (mathematics); Operations research; Management science; Artificial intelligence; Engineering; Mathematics","score_opus":0.07167254817495264,"score_gpt":0.3415309894101298,"score_spread":0.2698584412351772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116146752","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1437387,0.00081723486,0.84202296,0.00063378323,0.00004792922,0.00031635383,0.00026915167,0.0002278999,0.011925994],"genre_scores_gemma":[0.82937014,0.0006416691,0.16475622,0.0000762802,0.00004126318,0.00031863785,0.00030225562,0.000062718624,0.0044308156],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99705887,0.0009159936,0.00024101142,0.00048457374,0.00090400793,0.0003955398],"domain_scores_gemma":[0.9946956,0.003368463,0.0006242863,0.00025665926,0.0007638099,0.00029118604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038270515,0.0010997673,0.0009848024,0.0016437859,0.0009933652,0.0023744002,0.0013408484,0.0012922216,0.0021831694],"category_scores_gemma":[0.008237625,0.00076662254,0.0010882176,0.0014460545,0.0011328118,0.0024865605,0.0015812515,0.0011116618,0.00025378616],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023531268,0.00007480058,0.0012444576,0.0002278759,0.000070022805,0.0007433483,0.00057086913,0.908725,0.0036489274,0.042344958,0.0008290364,0.041285377],"study_design_scores_gemma":[0.000046299232,0.00013287332,0.00084717525,0.00005040696,0.000039679788,0.00014834218,0.0002264072,0.95672685,0.0020806699,0.038156897,0.0014992378,0.000045206532],"about_ca_topic_score_codex":0.0075662266,"about_ca_topic_score_gemma":0.005140667,"teacher_disagreement_score":0.0075662266,"about_ca_system_score_codex":0.0019250355,"about_ca_system_score_gemma":0.0015275148,"threshold_uncertainty_score":0.020239592},"labels":[],"label_agreement":null},{"id":"W2116218373","doi":"10.5555/2337223.2337429","title":"WorkItemExplorer: visualizing software development tasks using an interactive exploration environment","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Data exploration; Task (project management); Human–computer interaction; Ask price; Software; Task management; Software development; Task analysis; Software engineering; Visualization; Systems engineering; Artificial intelligence; Engineering","score_opus":0.09802627681252768,"score_gpt":0.3253566351947563,"score_spread":0.22733035838222865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116218373","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11082185,0.0008808823,0.8031943,0.0013330934,0.00018041028,0.00051151036,0.004315729,0.059738014,0.019024178],"genre_scores_gemma":[0.32363257,0.00076888234,0.66066206,0.00037838443,0.00007211105,0.00071663887,0.003465806,0.0033763428,0.006927205],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993755,0.00026358874,0.000035411253,0.000087279856,0.0001574711,0.000080747624],"domain_scores_gemma":[0.9972516,0.0017588034,0.00013298326,0.00037805125,0.0001834317,0.00029513278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013652958,0.0011908209,0.00043256674,0.0013404174,0.0005574912,0.002820127,0.0013615504,0.0009966362,0.011237787],"category_scores_gemma":[0.0033691896,0.00054675544,0.0008343899,0.0007607268,0.00052289845,0.0016899969,0.0037676198,0.0010095643,0.001327581],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0047137546,0.000970993,0.019537704,0.0024945764,0.00050652894,0.0028114477,0.02572348,0.034560747,0.13697276,0.02127429,0.13388559,0.6165481],"study_design_scores_gemma":[0.0009521788,0.0013947793,0.023963805,0.0012468512,0.00034955452,0.002395039,0.007390707,0.27858514,0.14015874,0.028566658,0.5141646,0.0008319653],"about_ca_topic_score_codex":0.0021089078,"about_ca_topic_score_gemma":0.0039089113,"teacher_disagreement_score":0.011237787,"about_ca_system_score_codex":0.00030212395,"about_ca_system_score_gemma":0.00060470303,"threshold_uncertainty_score":0.0375942},"labels":[],"label_agreement":null},{"id":"W2116219896","doi":"10.1145/1985793.1985955","title":"A software behaviour analysis framework based on the human perception systems (NIER track)","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Tracing; Variety (cybernetics); Software system; Process (computing); Software; Track (disk drive); Perception; Software development; Software engineering; Artificial intelligence; Programming language; Operating system","score_opus":0.04371383065218107,"score_gpt":0.27996268424954046,"score_spread":0.2362488535973594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116219896","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00076993403,0.00019744183,0.9924688,0.0002648269,0.000026704696,0.00017621335,0.00013184511,0.002178991,0.0037853864],"genre_scores_gemma":[0.0338795,0.00041943914,0.95978147,0.000162683,0.000057577403,0.00036536838,0.00038984374,0.00039717424,0.00454689],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99712354,0.0008210953,0.00015761513,0.00070583174,0.0010104057,0.00018152843],"domain_scores_gemma":[0.9982191,0.00074542296,0.00016014108,0.00041427396,0.0003293788,0.00013165512],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003523082,0.0015602591,0.0011818351,0.0041170395,0.0014570097,0.004028108,0.0031261723,0.0018152404,0.00842558],"category_scores_gemma":[0.005631234,0.0010484933,0.0033719158,0.001571564,0.0037686045,0.0040578805,0.0026259543,0.0031289954,0.0037110539],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000919016,0.00023443783,0.002066086,0.00058808975,0.00018503584,0.00036594828,0.0018092522,0.026736727,0.011268302,0.70277053,0.010282569,0.24360114],"study_design_scores_gemma":[0.00005804809,0.00023991137,0.003299315,0.00038088058,0.0001510069,0.00055138976,0.000427625,0.31400552,0.008041905,0.5499433,0.12275768,0.0001432676],"about_ca_topic_score_codex":0.01651846,"about_ca_topic_score_gemma":0.011962467,"teacher_disagreement_score":0.01651846,"about_ca_system_score_codex":0.0018933861,"about_ca_system_score_gemma":0.0033924377,"threshold_uncertainty_score":0.032844663},"labels":[],"label_agreement":null},{"id":"W2116228345","doi":"10.1109/icsme.2014.107","title":"Reviewer Recommender of Pull-Requests in GitHub","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":116,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Recommender system; World Wide Web; Process (computing); Precision and recall; Information retrieval; Recall; Quality (philosophy); Software; Code (set theory); Crowds; Code review; Software development; Software quality; Computer security","score_opus":0.02409844039246137,"score_gpt":0.2869606255315675,"score_spread":0.2628621851391061,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116228345","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19571505,0.006210608,0.36560896,0.0044002985,0.0025841233,0.0037154567,0.028800065,0.3383017,0.0546638],"genre_scores_gemma":[0.37727627,0.0016633334,0.50017,0.0015585136,0.0007861033,0.0011714625,0.04563459,0.010709997,0.06102975],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959122,0.001249496,0.00030504738,0.000881808,0.0013641536,0.00028742134],"domain_scores_gemma":[0.9841489,0.004349247,0.00094857323,0.0037829622,0.0053462205,0.0014240884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043184385,0.003002711,0.002083991,0.005348314,0.0020371762,0.0029797193,0.002312007,0.002321429,0.009909065],"category_scores_gemma":[0.021572514,0.0012935097,0.0014071367,0.0032382126,0.00037199297,0.0033342787,0.0022747351,0.0015140336,0.017614553],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003604802,0.00096436986,0.048711184,0.002343288,0.0008205732,0.0024846639,0.00215475,0.014687114,0.03152167,0.0051955855,0.45045727,0.43705484],"study_design_scores_gemma":[0.00063679664,0.00084506476,0.022560716,0.00030365473,0.0006036843,0.0014678156,0.0010121907,0.6404495,0.04148768,0.008706192,0.28135008,0.0005765565],"about_ca_topic_score_codex":0.018470246,"about_ca_topic_score_gemma":0.043309268,"teacher_disagreement_score":0.018470246,"about_ca_system_score_codex":0.0010188703,"about_ca_system_score_gemma":0.002757486,"threshold_uncertainty_score":0.03672552},"labels":[],"label_agreement":null},{"id":"W2116255095","doi":"10.1109/icsm.2006.38","title":"Mining Software Repositories to Assist Developers and Support Managers","year":2006,"lang":"en","type":"article","venue":"Proceedings/Proceedings - Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Work (physics); Software; Plan (archaeology); Software development; Software engineering; Software project management; Software evolution; Software analytics; Open source software; World Wide Web; Software peer review; Data science; Social software engineering; Software construction; Engineering","score_opus":0.01846161005789775,"score_gpt":0.24537692021052368,"score_spread":0.22691531015262592,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116255095","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4488024,0.0031496303,0.4993849,0.009202078,0.00018416326,0.0021726748,0.0075756856,0.019637544,0.0098909885],"genre_scores_gemma":[0.4899685,0.0013523657,0.4892362,0.00029451592,0.000111649046,0.000630284,0.014814657,0.0003916726,0.0032000165],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9950682,0.0017415796,0.0006726568,0.0006796541,0.0015713047,0.00026660706],"domain_scores_gemma":[0.967076,0.01425309,0.006678731,0.0046936036,0.006638633,0.00065986085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007399403,0.0010803029,0.0012519724,0.014545644,0.0012854266,0.004703776,0.0026547683,0.0016528055,0.001864187],"category_scores_gemma":[0.036949404,0.00082404737,0.0008010609,0.008297047,0.00044334796,0.005944771,0.0023200014,0.001159477,0.001736406],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024129341,0.0014256346,0.12432944,0.001023826,0.000288214,0.00044177918,0.0022218595,0.012102326,0.007907102,0.005647821,0.031793833,0.81257695],"study_design_scores_gemma":[0.0005877316,0.0012965337,0.12707762,0.0017258321,0.0011388771,0.0022772795,0.017221788,0.583669,0.06306374,0.07139961,0.13019665,0.0003453452],"about_ca_topic_score_codex":0.0051040393,"about_ca_topic_score_gemma":0.00969384,"teacher_disagreement_score":0.014545644,"about_ca_system_score_codex":0.00088195025,"about_ca_system_score_gemma":0.0037004538,"threshold_uncertainty_score":0.039132297},"labels":[],"label_agreement":null},{"id":"W2116295136","doi":"10.1109/metric.2003.1232451","title":"Metrology, measurement and metrics in software engineering","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software measurement; Metrology; Computer science; Software; Software engineering; Domain (mathematical analysis); Software development; Software Engineering Process Group; Position paper; Social software engineering; Systems engineering; Perspective (graphical); Software construction; Engineering; Artificial intelligence","score_opus":0.026587276611645866,"score_gpt":0.2377850616378372,"score_spread":0.21119778502619133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116295136","genre_codex":"review","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008088455,0.46913928,0.32127118,0.04772517,0.0037944799,0.00017459015,0.00025954857,0.00053173845,0.14901555],"genre_scores_gemma":[0.45780608,0.24764207,0.24024728,0.010355498,0.0089860875,0.0010370855,0.00044516337,0.0004566908,0.03302405],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9861373,0.0076529034,0.00094443583,0.0012182493,0.0036976475,0.0003494741],"domain_scores_gemma":[0.98349863,0.012648179,0.0012733935,0.0009758025,0.0012765215,0.00032753937],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0079794405,0.0015927152,0.0017728768,0.0076469197,0.0026253175,0.008771805,0.0018115083,0.005633774,0.0038322704],"category_scores_gemma":[0.02240195,0.00076726277,0.00067221187,0.013517179,0.026081815,0.017107874,0.004389581,0.0062477416,0.0013687622],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009218586,0.000010518372,0.0002471702,0.00031986507,0.000007705371,0.00006527721,0.00096642395,0.00097609643,0.000100655,0.96479005,0.004306528,0.028200632],"study_design_scores_gemma":[0.000003132267,0.000021024163,0.00037046187,0.0004028834,0.0000047073927,0.0001634754,0.0004042366,0.0010188542,0.00009014264,0.912543,0.084955595,0.000022513963],"about_ca_topic_score_codex":0.0047111595,"about_ca_topic_score_gemma":0.0032103015,"teacher_disagreement_score":0.99202055,"about_ca_system_score_codex":0.0051416475,"about_ca_system_score_gemma":0.0037825617,"threshold_uncertainty_score":0.04219979},"labels":[],"label_agreement":null},{"id":"W2116357421","doi":"10.1145/13487689.13487691","title":"Topology analysis of software dependencies","year":2008,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":106,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Intuition; Source code; Task (project management); Static program analysis; Set (abstract data type); Dependency (UML); Software; Program analysis; Fuzzy logic; Programming language; Software engineering; Data mining; Theoretical computer science; Software development; Artificial intelligence; Systems engineering","score_opus":0.09963517125290641,"score_gpt":0.32526599820654656,"score_spread":0.22563082695364015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116357421","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.355107,0.00054016546,0.63205165,0.00019798729,0.000027593123,0.0002038154,0.0011590231,0.0016435302,0.009069142],"genre_scores_gemma":[0.8431698,0.00031646524,0.15314306,0.000022361808,0.000028477945,0.00015220813,0.0013061303,0.00018501827,0.0016764305],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988888,0.00023367487,0.00007411212,0.00017309077,0.0005214809,0.00010899149],"domain_scores_gemma":[0.9924523,0.004070837,0.00097363273,0.0005325613,0.0016987245,0.0002718887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00080035697,0.000475609,0.0004778739,0.009035432,0.000898425,0.0012290154,0.0005961668,0.0006912605,0.0031873777],"category_scores_gemma":[0.0101045575,0.00043057784,0.00074992725,0.0029745277,0.0006364359,0.0018552567,0.0008162188,0.00056848244,0.00050692505],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006002327,0.00020008784,0.06587483,0.00070384855,0.00019844755,0.0010577643,0.0022054375,0.3435285,0.059789963,0.08554745,0.0057346174,0.43455884],"study_design_scores_gemma":[0.00001784388,0.00013716359,0.027332252,0.000054355078,0.00007540002,0.00066380535,0.000515606,0.89809406,0.013262604,0.052761763,0.0070227482,0.00006246495],"about_ca_topic_score_codex":0.003563928,"about_ca_topic_score_gemma":0.0037618023,"teacher_disagreement_score":0.009035432,"about_ca_system_score_codex":0.0008121468,"about_ca_system_score_gemma":0.00073932175,"threshold_uncertainty_score":0.010662854},"labels":[],"label_agreement":null},{"id":"W2116492147","doi":"10.1109/wpc.2003.1199186","title":"Design recovery of a two level system","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Programming language; Interpreter; Compiler; Grammar; Natural language processing; Linguistics","score_opus":0.06655882008245029,"score_gpt":0.27997506793743937,"score_spread":0.21341624785498908,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116492147","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07371016,0.00012340368,0.9011014,0.0012418403,0.00010838592,0.0003255619,0.00024802925,0.011346849,0.011794389],"genre_scores_gemma":[0.4952271,0.0001601762,0.47534463,0.0006354703,0.00003145955,0.00042885495,0.00089509424,0.0012073583,0.026069822],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99747735,0.0005118744,0.00016774138,0.00032851062,0.0011936047,0.00032090925],"domain_scores_gemma":[0.9963542,0.0006548776,0.0002190722,0.0019948652,0.0006673991,0.00010963124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019090425,0.0005004327,0.0005300243,0.00086170697,0.0008702393,0.0022014051,0.0011350515,0.0017664412,0.0058408994],"category_scores_gemma":[0.0064667393,0.00069095916,0.0010380682,0.0005504571,0.0013179497,0.0021796394,0.0024416267,0.0023379794,0.0026196737],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073813915,0.00063978654,0.0075713936,0.00077924115,0.00013371295,0.0036531286,0.005343624,0.15521581,0.1654907,0.31217432,0.021298453,0.32696167],"study_design_scores_gemma":[0.0002386646,0.0006906183,0.002158907,0.0001705275,0.00015184049,0.0015226044,0.0007232598,0.5939541,0.14404814,0.10652198,0.14967933,0.00013999506],"about_ca_topic_score_codex":0.0028744119,"about_ca_topic_score_gemma":0.0022783666,"teacher_disagreement_score":0.0058408994,"about_ca_system_score_codex":0.0017337649,"about_ca_system_score_gemma":0.003360286,"threshold_uncertainty_score":0.019539773},"labels":[],"label_agreement":null},{"id":"W2116595917","doi":"10.1109/csmr.2008.4493313","title":"On the Maintainability of Aspect-Oriented Software: A Concern-Oriented Measurement Framework","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":74,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Coordenação de Aperfeiçoamento de Pessoal de Nível Superior; European Commission","keywords":"Maintainability; Terminology; Aspect-oriented programming; Computer science; Generality; Software engineering; Process (computing); Modularity (biology); Software maintenance; Risk analysis (engineering); Software; Process management; Management science; Software development; Programming language; Engineering","score_opus":0.04098637484080733,"score_gpt":0.2679368122798093,"score_spread":0.22695043743900198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116595917","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028470207,0.002165216,0.96195364,0.0011397537,0.000043604126,0.00027711564,0.00016466416,0.00035219593,0.0054336977],"genre_scores_gemma":[0.3235149,0.0019646806,0.67235386,0.0002322736,0.000108344,0.0010187625,0.00030088035,0.000103668885,0.0004026416],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9799692,0.009216631,0.0025653332,0.0017802169,0.005782603,0.0006859494],"domain_scores_gemma":[0.95052326,0.03138515,0.006366083,0.0037884484,0.007016161,0.0009209882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026241481,0.002071024,0.001338873,0.018904535,0.0020904795,0.006783102,0.0022762627,0.0023903558,0.00064919906],"category_scores_gemma":[0.04936749,0.00075697043,0.0016064865,0.009897811,0.007818985,0.009744622,0.004496191,0.002606389,0.00025352908],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007229867,0.00026120013,0.020827781,0.0008434924,0.00016962756,0.00026678742,0.004217082,0.02922058,0.00739131,0.78986007,0.0014231349,0.14544661],"study_design_scores_gemma":[0.000048118127,0.0009556204,0.039281122,0.0018736315,0.00030680365,0.0015902532,0.0031518606,0.1963716,0.00849877,0.71641475,0.031157006,0.00035038716],"about_ca_topic_score_codex":0.0048021604,"about_ca_topic_score_gemma":0.0027807017,"teacher_disagreement_score":0.026241481,"about_ca_system_score_codex":0.0034636636,"about_ca_system_score_gemma":0.0039034162,"threshold_uncertainty_score":0.13877988},"labels":[],"label_agreement":null},{"id":"W2116759993","doi":"10.1016/j.jss.2005.05.001","title":"Automated impact analysis of UML models","year":2005,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":119,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Goddard Space Flight Center","keywords":"Unified Modeling Language; Applications of UML; Computer science; UML tool; Class diagram; Object Constraint Language; Data mining; Software engineering; Programming language; Software","score_opus":0.020332780479825893,"score_gpt":0.2926373504047893,"score_spread":0.2723045699249634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116759993","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39730346,0.0008813635,0.57581,0.00060316914,0.00010864595,0.00047900146,0.0011881638,0.012572355,0.011053818],"genre_scores_gemma":[0.85600734,0.0002847825,0.14000365,0.000045218323,0.000041614825,0.00013622663,0.0017283496,0.00050417683,0.0012486319],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935487,0.0016660965,0.0003004225,0.00030073925,0.0038589565,0.00032524057],"domain_scores_gemma":[0.9783448,0.013974209,0.0015346303,0.0025744096,0.0033818136,0.00019015037],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028544934,0.0010043336,0.0008702538,0.0055969693,0.0006845317,0.0017554021,0.0014728773,0.00075591507,0.0029864586],"category_scores_gemma":[0.02417193,0.00071860105,0.0014523371,0.002091147,0.00053645554,0.0029605872,0.0019795566,0.00087978016,0.00053875305],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008231121,0.0006321765,0.025796587,0.00063384883,0.0003439478,0.00048843503,0.00052725966,0.29156885,0.037152935,0.040364858,0.0052493364,0.5964186],"study_design_scores_gemma":[0.000033733497,0.0000729151,0.0034360094,0.000030403886,0.00009619553,0.00007003307,0.00007952621,0.9714253,0.010783586,0.012017152,0.0019373376,0.000017748658],"about_ca_topic_score_codex":0.0068818484,"about_ca_topic_score_gemma":0.0098992465,"teacher_disagreement_score":0.0068818484,"about_ca_system_score_codex":0.0015862188,"about_ca_system_score_gemma":0.0017019056,"threshold_uncertainty_score":0.015096188},"labels":[],"label_agreement":null},{"id":"W2116871579","doi":"10.1007/3-540-44870-5_16","title":"Tool Support for Complex Refactoring to Design Patterns","year":2003,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Code refactoring; Maintainability; Computer science; Software design pattern; Agile software development; Software engineering; Structural pattern; Software design; Design pattern; Software system; Architectural pattern; Software; Software development; Programming language","score_opus":0.06321337205430425,"score_gpt":0.29854956235159047,"score_spread":0.23533619029728622,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116871579","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008622602,0.0002796071,0.92923576,0.00016356005,0.000056623034,0.00009662106,0.00037678104,0.05754214,0.0036263617],"genre_scores_gemma":[0.10426567,0.00063397794,0.86937076,0.00025488724,0.00006310075,0.00025969863,0.0030305395,0.011855716,0.010265693],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982352,0.00032236482,0.00031922598,0.00029389944,0.0007044299,0.00012470288],"domain_scores_gemma":[0.98879963,0.0067481343,0.0005169976,0.002993894,0.0007658183,0.00017555185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002493611,0.0015101151,0.0010485466,0.0022327981,0.00047850108,0.0023653733,0.003138668,0.00152756,0.013743847],"category_scores_gemma":[0.012820098,0.0014960999,0.0014546592,0.0020240673,0.0005846624,0.0038612464,0.002682372,0.0021628693,0.005001656],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002562484,0.00026331772,0.0019339942,0.00095999683,0.00012220714,0.00091486634,0.0012742063,0.008080241,0.03391715,0.027167387,0.032779273,0.89233106],"study_design_scores_gemma":[0.0006660435,0.000353112,0.0026761128,0.0011250498,0.0005038448,0.0042182812,0.00039751158,0.38150617,0.18488835,0.11349617,0.30986252,0.00030685661],"about_ca_topic_score_codex":0.00082005316,"about_ca_topic_score_gemma":0.0013049472,"teacher_disagreement_score":0.013743847,"about_ca_system_score_codex":0.00042320765,"about_ca_system_score_gemma":0.000610211,"threshold_uncertainty_score":0.04597777},"labels":[],"label_agreement":null},{"id":"W2117118977","doi":"10.1109/csmr.2009.12","title":"A Case Study of Source Code Evolution","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Software quality; Computer science; Source code; Fault (geology); Software engineering; Code (set theory); Software; Software development; Quality (philosophy); Process (computing); Reliability engineering; Resource (disambiguation); Software evolution; Product (mathematics); Real-time computing; Software construction; Engineering; Operating system; Programming language","score_opus":0.02605842075428154,"score_gpt":0.28957574738605213,"score_spread":0.2635173266317706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117118977","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98164064,0.00029191,0.013711996,0.00057887414,0.000023482608,0.00015292066,0.0003788938,0.0002862951,0.0029348785],"genre_scores_gemma":[0.97474414,0.0001850187,0.02229936,0.000096583484,0.0000148481295,0.00007898702,0.0006331753,0.00009808312,0.0018499093],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99607193,0.0014776073,0.0002593691,0.0005070097,0.0014222104,0.00026201783],"domain_scores_gemma":[0.95748526,0.030588562,0.0026675884,0.003271774,0.004917401,0.0010693643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034390613,0.0004294614,0.0003324185,0.003095749,0.0018644859,0.0009762769,0.0012549417,0.0018363608,0.00091988355],"category_scores_gemma":[0.022961384,0.00036048846,0.0005167439,0.0040575373,0.001079011,0.0011731933,0.00095651444,0.0011355316,0.00022749431],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011851705,0.00260409,0.39774054,0.0010995155,0.0003242493,0.0757536,0.042121567,0.080875464,0.040802278,0.012429435,0.011738186,0.33332592],"study_design_scores_gemma":[0.00039606815,0.0026490565,0.46724242,0.0003471914,0.0003296139,0.049918503,0.022944188,0.29272246,0.057645872,0.010851781,0.094639435,0.00031343105],"about_ca_topic_score_codex":0.013082748,"about_ca_topic_score_gemma":0.020988781,"teacher_disagreement_score":0.013082748,"about_ca_system_score_codex":0.0013045209,"about_ca_system_score_gemma":0.0010684488,"threshold_uncertainty_score":0.026013196},"labels":[],"label_agreement":null},{"id":"W2117508687","doi":"10.1145/1858996.1859050","title":"An experience report on scaling tools for mining software repositories using MapReduce","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Scalability; Software engineering; Software; Field (mathematics); Search-based software engineering; Data science; Scale (ratio); Heuristic; Software development; Software construction; Database; Operating system; Artificial intelligence","score_opus":0.051133250636763404,"score_gpt":0.34706361563746463,"score_spread":0.29593036500070125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117508687","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46784866,0.0051874104,0.39852893,0.010819748,0.00068407535,0.002244608,0.003385466,0.027067201,0.08423393],"genre_scores_gemma":[0.44346893,0.0046198214,0.521744,0.0014746205,0.00029076973,0.0005704245,0.0057900357,0.0024689923,0.0195724],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9947594,0.0011282612,0.00033736104,0.00064123207,0.0028029026,0.00033091634],"domain_scores_gemma":[0.97987694,0.0077595767,0.0005760006,0.004404062,0.0060691372,0.001314269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009722134,0.0008041553,0.00048583146,0.0033731884,0.001401914,0.0028277512,0.0017472211,0.0008441013,0.0027929367],"category_scores_gemma":[0.024604227,0.00055423187,0.0009186508,0.0056041954,0.00091252464,0.0046883957,0.0017800404,0.0014473778,0.001739386],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026653547,0.0014560693,0.022050923,0.00071116636,0.00020701835,0.0009907254,0.0075751175,0.006349507,0.026082717,0.0058758496,0.056786623,0.8716477],"study_design_scores_gemma":[0.0002826301,0.0026549958,0.06443459,0.0006235136,0.00044220738,0.0044879345,0.011007334,0.06958931,0.08997641,0.013652262,0.742471,0.00037783364],"about_ca_topic_score_codex":0.0075544273,"about_ca_topic_score_gemma":0.008213371,"teacher_disagreement_score":0.009722134,"about_ca_system_score_codex":0.0007542583,"about_ca_system_score_gemma":0.0024170408,"threshold_uncertainty_score":0.05141616},"labels":[],"label_agreement":null},{"id":"W2117757207","doi":"10.5555/2664446.2664448","title":"Towards improving bug tracking systems with game mechanisms","year":2012,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Agile software development; Tracking system; Quality (philosophy); Tracking (education); Reputation; Software engineering; Computer security; Human–computer interaction; Artificial intelligence","score_opus":0.016023266705648975,"score_gpt":0.2417133949857941,"score_spread":0.22569012828014515,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117757207","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2284522,0.0013741672,0.74481153,0.0039394903,0.0001740057,0.0017247443,0.00018156979,0.014493053,0.004849255],"genre_scores_gemma":[0.41617188,0.0003601163,0.5800859,0.00043684753,0.00006574838,0.0006625094,0.00024971686,0.00042720762,0.0015401285],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9741469,0.014814288,0.0023566869,0.002788761,0.004832205,0.0010611378],"domain_scores_gemma":[0.8492488,0.07648979,0.020185366,0.027919468,0.021555187,0.0046014027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04116054,0.0018904903,0.0014686637,0.008433836,0.0017330033,0.0074062347,0.004907247,0.0023980138,0.002657325],"category_scores_gemma":[0.1769607,0.0015512309,0.0014540899,0.0035446703,0.0025819307,0.012927097,0.0061343783,0.0029240085,0.0011553917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007480231,0.0030292596,0.07757057,0.0018360469,0.00048863294,0.0003031157,0.0055301613,0.03334528,0.037492372,0.036396023,0.0077032195,0.79555726],"study_design_scores_gemma":[0.0009737936,0.004037373,0.03941377,0.0011048996,0.00087807275,0.0010390257,0.003641111,0.7403758,0.03698729,0.118830785,0.05214976,0.00056831114],"about_ca_topic_score_codex":0.002730855,"about_ca_topic_score_gemma":0.0038160179,"teacher_disagreement_score":0.04116054,"about_ca_system_score_codex":0.002179718,"about_ca_system_score_gemma":0.00435208,"threshold_uncertainty_score":0.21768034},"labels":[],"label_agreement":null},{"id":"W2117852863","doi":"10.1109/icse-companion.2009.5070996","title":"Can peer code reviews be exploited for later information needs?","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Code (set theory); Programming language","score_opus":0.04468858656087239,"score_gpt":0.29796818528718233,"score_spread":0.25327959872630995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117852863","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16881225,0.00965715,0.13949041,0.32535398,0.005466055,0.0023079428,0.0005371637,0.008203643,0.3401714],"genre_scores_gemma":[0.8059135,0.005058554,0.07056716,0.019507486,0.0039886436,0.0014562667,0.0006541531,0.0020294841,0.09082469],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.90499383,0.055236135,0.00413434,0.006119101,0.024786886,0.0047296807],"domain_scores_gemma":[0.4991472,0.23647983,0.041835256,0.08735515,0.12167627,0.013506219],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07398676,0.0008642932,0.0014870011,0.007687695,0.004696885,0.013716887,0.0040887357,0.0063022976,0.028875131],"category_scores_gemma":[0.49527937,0.0012472199,0.001218978,0.0052522304,0.0043858206,0.048352063,0.00810249,0.0036478133,0.019029435],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034779153,0.00044828487,0.021849511,0.001174017,0.00008930134,0.0010590997,0.023052096,0.00033197514,0.003428964,0.031870324,0.100127265,0.81622136],"study_design_scores_gemma":[0.00041516195,0.0012859389,0.049033273,0.0043912553,0.00031564286,0.005049428,0.04767028,0.0071131256,0.008815027,0.14358576,0.73177516,0.0005499845],"about_ca_topic_score_codex":0.0043124924,"about_ca_topic_score_gemma":0.0045659794,"teacher_disagreement_score":0.92601323,"about_ca_system_score_codex":0.0032426312,"about_ca_system_score_gemma":0.016477762,"threshold_uncertainty_score":0.39128405},"labels":[],"label_agreement":null},{"id":"W2117980885","doi":"10.1109/tools.1999.787538","title":"A new metrics set for evaluating testing efforts for object-oriented programs","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Object-oriented programming; Set (abstract data type); Software metric; Software; Programming language; Inheritance (genetic algorithm); Software engineering; Object (grammar); Software development; Software construction; Artificial intelligence","score_opus":0.11569806912040391,"score_gpt":0.3668609365252554,"score_spread":0.2511628674048515,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117980885","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09777587,0.0043544406,0.8796444,0.0009118111,0.0003859883,0.001711163,0.0027308865,0.0019938734,0.010491625],"genre_scores_gemma":[0.33101174,0.0011119989,0.6586105,0.00019602187,0.00014598038,0.0025720543,0.0043556658,0.00039448452,0.0016015868],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.95742965,0.011910795,0.0062454846,0.0015839827,0.022168,0.0006621113],"domain_scores_gemma":[0.88622326,0.05972153,0.012440047,0.009424004,0.029483467,0.002707779],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017625002,0.0024233207,0.0021356805,0.020373825,0.0012855,0.0035929077,0.0022514767,0.0016658455,0.001676659],"category_scores_gemma":[0.11900958,0.0004310573,0.001816292,0.009983226,0.0014085739,0.008072813,0.003134006,0.0016978686,0.00050645607],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064233126,0.00073331804,0.081189394,0.0019390351,0.0010289961,0.0002772526,0.0017005806,0.060662523,0.01481185,0.05336221,0.012798213,0.77085423],"study_design_scores_gemma":[0.0004878641,0.011975263,0.16915126,0.0020017275,0.0022654352,0.0032555477,0.0027468791,0.48455095,0.04608747,0.17836536,0.097976826,0.0011354707],"about_ca_topic_score_codex":0.0021895315,"about_ca_topic_score_gemma":0.0022340466,"teacher_disagreement_score":0.020373825,"about_ca_system_score_codex":0.0031896331,"about_ca_system_score_gemma":0.0032027238,"threshold_uncertainty_score":0.093211055},"labels":[],"label_agreement":null},{"id":"W2118397610","doi":"10.1109/icse.2009.5070559","title":"ConcernLines: A timeline view of co-occurring concerns","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Timeline; Computer science; Software evolution; Software engineering; Process (computing); Software; Software system; Software construction; Programming language; History","score_opus":0.03853328266429952,"score_gpt":0.34346241191083454,"score_spread":0.30492912924653504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118397610","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018523965,0.001119283,0.93407935,0.0012254539,0.00034162146,0.0001833618,0.0035503926,0.022591606,0.01838503],"genre_scores_gemma":[0.1776396,0.0017078662,0.79957634,0.00039334025,0.00023463048,0.0003969354,0.00449669,0.004725368,0.010829289],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992488,0.00022679407,0.00006581171,0.00013384501,0.0002545235,0.00007018466],"domain_scores_gemma":[0.9959092,0.0021477812,0.00055355194,0.00047055728,0.00052623876,0.00039254114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013190504,0.0013901407,0.00046483675,0.0028682987,0.0009001816,0.0043277876,0.0014584132,0.0014192206,0.01135932],"category_scores_gemma":[0.006016215,0.00074575644,0.0009031417,0.0019437965,0.00096159807,0.008082788,0.0024988793,0.0027553889,0.0017048279],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011522867,0.0003042051,0.011881492,0.0018829874,0.00029162472,0.0025123267,0.034037814,0.023369126,0.044476837,0.402957,0.13602409,0.34111023],"study_design_scores_gemma":[0.00012771669,0.00021274558,0.0048645013,0.00035918987,0.00016629905,0.0018781717,0.004042767,0.075254075,0.01449286,0.12849551,0.7699069,0.00019934164],"about_ca_topic_score_codex":0.008060405,"about_ca_topic_score_gemma":0.011385295,"teacher_disagreement_score":0.01135932,"about_ca_system_score_codex":0.0007337655,"about_ca_system_score_gemma":0.0012813848,"threshold_uncertainty_score":0.038000762},"labels":[],"label_agreement":null},{"id":"W2118466151","doi":"10.1109/re.2005.58","title":"Requirements engineering and the creative process in the video game industry","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":142,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Video game; Computer science; Process (computing); Game Developer; Field (mathematics); Domain (mathematical analysis); Documentation; Game art design; Video game development; Game design; Work in process; Multimedia; Engineering; Operations management","score_opus":0.021867665705451816,"score_gpt":0.2931609048746015,"score_spread":0.2712932391691497,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118466151","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.745032,0.007654112,0.12674244,0.013717694,0.000073266114,0.0004694179,0.00003953368,0.0002143208,0.10605724],"genre_scores_gemma":[0.96541744,0.0022645113,0.029229859,0.00041700754,0.000020033951,0.000117081276,0.000029649938,0.00003302583,0.0024713678],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9689003,0.024701122,0.0011269986,0.000643626,0.003668437,0.0009595461],"domain_scores_gemma":[0.8762862,0.1130324,0.0047794026,0.0019133395,0.0028995704,0.0010890875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021714292,0.0005669908,0.00028621536,0.0024902187,0.0020017494,0.005434482,0.0011219832,0.002396871,0.0016600499],"category_scores_gemma":[0.06707923,0.00063388224,0.00041319418,0.001969655,0.008594001,0.0047512273,0.0023156009,0.002078092,0.00038778447],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004545878,0.0011243916,0.023484372,0.0017831004,0.0000741196,0.005265414,0.18407974,0.022912042,0.008510905,0.33612558,0.004051052,0.4121347],"study_design_scores_gemma":[0.0004281723,0.001369685,0.049636412,0.0039470657,0.0000879637,0.008589961,0.19348963,0.077707216,0.0136695905,0.49861056,0.15195327,0.00051046856],"about_ca_topic_score_codex":0.003437988,"about_ca_topic_score_gemma":0.0038652448,"teacher_disagreement_score":0.021714292,"about_ca_system_score_codex":0.0028570364,"about_ca_system_score_gemma":0.0031469963,"threshold_uncertainty_score":0.11483747},"labels":[],"label_agreement":null},{"id":"W2118604270","doi":"10.1109/tse.2012.28","title":"The Effects of Test-Driven Development on External Quality and Productivity: A Meta-Analysis","year":2013,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":125,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Productivity; Moderation; Computer science; Meta-analysis; Quality (philosophy); Task (project management); Quality management; Statistics; Operations management; Mathematics; Engineering; Economics","score_opus":0.03296412813828898,"score_gpt":0.27103694364331116,"score_spread":0.23807281550502218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118604270","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019800285,0.9763016,0.0018154972,0.00026224885,0.00024478597,0.00045036193,0.00065277977,0.000055661727,0.00041671557],"genre_scores_gemma":[0.44631258,0.54264915,0.0062356815,0.0008380781,0.00030259232,0.0017728641,0.0012973384,0.00010354056,0.00048817086],"study_design_codex":"meta_analysis","study_design_gemma":"meta_analysis","domain_scores_codex":[0.98186284,0.009493257,0.004717534,0.001571752,0.0020023226,0.00035228883],"domain_scores_gemma":[0.94297147,0.0462486,0.005831053,0.0018941294,0.002543194,0.00051157904],"candidate_categories":["metaresearch","metaepi_broad"],"consensus_categories":[],"category_scores_codex":[0.025042087,0.0029373362,0.014091789,0.0074247764,0.0007245502,0.0038367494,0.0022295106,0.0019016683,0.0021358645],"category_scores_gemma":[0.06032579,0.0015542982,0.047083456,0.006997239,0.0008938088,0.0018648756,0.00178725,0.0022395328,0.00021735948],"study_design_candidate":"meta_analysis","study_design_consensus":"meta_analysis","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002152672,0.000038809747,0.0046465355,0.13113195,0.8508765,0.000107422245,0.00009128379,0.00040926068,0.00028661004,0.00009181698,0.00024871054,0.00991849],"study_design_scores_gemma":[0.00047421287,0.00024292644,0.0034468102,0.008394125,0.98627937,0.000071658214,0.00003344129,0.00010144776,0.00022172074,0.00011068706,0.0006091558,0.000014430292],"about_ca_topic_score_codex":0.004917013,"about_ca_topic_score_gemma":0.009458418,"teacher_disagreement_score":0.9859082,"about_ca_system_score_codex":0.0031145776,"about_ca_system_score_gemma":0.0027136186,"threshold_uncertainty_score":0.13243681},"labels":[],"label_agreement":null},{"id":"W2118655104","doi":"10.1109/icst.2012.106","title":"@tComment: Testing Javadoc Comments to Detect Comment-Code Inconsistencies","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":184,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Java; Programming language; Component (thermodynamics); Code (set theory); Set (abstract data type); Source code; Null (SQL); Reading (process); Software bug; Software; Random testing; Source lines of code; Information retrieval; Software engineering; Data mining; Test case; Machine learning","score_opus":0.0672872285367884,"score_gpt":0.30009323688933365,"score_spread":0.23280600835254525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118655104","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6721111,0.00024016459,0.2476621,0.00048919476,0.0002390244,0.0010617463,0.003915236,0.07080216,0.003479236],"genre_scores_gemma":[0.71056277,0.00008091912,0.27492204,0.0003274519,0.000058950816,0.00065637473,0.005046901,0.0058837505,0.0024608504],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9801184,0.006937957,0.0023251225,0.0036804664,0.006051102,0.00088691537],"domain_scores_gemma":[0.804497,0.11911526,0.01983738,0.024498038,0.029486591,0.0025656964],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01165607,0.0016012287,0.00094618765,0.0031281568,0.00084463676,0.0014276168,0.0031035494,0.0022752837,0.0032165593],"category_scores_gemma":[0.08475823,0.0010835971,0.0012328388,0.0017620543,0.0013643313,0.003484401,0.0021547864,0.0012853863,0.0015639367],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003105761,0.0026345134,0.344397,0.0023940664,0.0006930451,0.0030414553,0.007010233,0.023346538,0.1322792,0.0055488176,0.032591745,0.44295755],"study_design_scores_gemma":[0.00093032146,0.0037002086,0.122450866,0.00041670998,0.00043720278,0.0035181795,0.002235695,0.49305338,0.3336775,0.0055342773,0.033420455,0.0006252525],"about_ca_topic_score_codex":0.0045415433,"about_ca_topic_score_gemma":0.0056977407,"teacher_disagreement_score":0.01165607,"about_ca_system_score_codex":0.00089071295,"about_ca_system_score_gemma":0.0024778754,"threshold_uncertainty_score":0.061643958},"labels":[],"label_agreement":null},{"id":"W2118657783","doi":"10.1109/apsec.2000.896734","title":"Predicting class libraries interface evolution: an investigation into machine learning approaches","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Machine learning; Class (philosophy); Interface (matter); Artificial intelligence; Compatibility (geochemistry); Fuzzy logic; Set (abstract data type); Software engineering; Programming language; Engineering; Operating system","score_opus":0.05361013478208405,"score_gpt":0.23643412476935308,"score_spread":0.18282398998726904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118657783","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29839647,0.0018003664,0.695126,0.00083158986,0.00003335085,0.00013636186,0.00016257107,0.00089899026,0.002614308],"genre_scores_gemma":[0.8734163,0.0006206743,0.12494538,0.000084634536,0.000045735567,0.00008235528,0.00020647721,0.000040695322,0.0005577679],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99799335,0.0009327784,0.00013831524,0.00029824907,0.0005413086,0.00009602257],"domain_scores_gemma":[0.96838623,0.02794016,0.0010051512,0.0005238531,0.0019374039,0.0002071485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058077304,0.0008407031,0.00075934414,0.0033999213,0.0005156,0.0012796957,0.0009866388,0.001124841,0.00081544253],"category_scores_gemma":[0.02299649,0.00022236357,0.0006053763,0.0016968319,0.0005254624,0.0016880867,0.00045648712,0.0011095855,0.00030294046],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028276164,0.00045556863,0.07506247,0.00024346143,0.0002040688,0.000095997304,0.00031205663,0.2847154,0.0037294917,0.0054735714,0.0010002281,0.62842494],"study_design_scores_gemma":[0.000008905995,0.00010640084,0.005513705,0.000027615324,0.000031055435,0.000034624758,0.000060434813,0.98815954,0.0018887907,0.0038508093,0.0003042181,0.00001383527],"about_ca_topic_score_codex":0.0034500232,"about_ca_topic_score_gemma":0.0029310517,"teacher_disagreement_score":0.0058077304,"about_ca_system_score_codex":0.0007775901,"about_ca_system_score_gemma":0.00066428777,"threshold_uncertainty_score":0.030714571},"labels":[],"label_agreement":null},{"id":"W2118771704","doi":"10.1109/wpc.1996.501130","title":"A formal architectural design patterns-based approach to software understanding","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Architectural pattern; Abstraction; Formal methods; Software engineering; Programming language; Software design pattern; Software design; Formal specification; Formal verification; Software architecture; Code generation; Software; Software development; Theoretical computer science","score_opus":0.10888694081381078,"score_gpt":0.2551833144629549,"score_spread":0.1462963736491441,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118771704","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008055945,0.00022256417,0.99473673,0.0009829428,0.00003051347,0.00007586242,0.000052104187,0.00014546048,0.0029483305],"genre_scores_gemma":[0.024971778,0.00062744855,0.97141176,0.00025190707,0.000052370946,0.0003560934,0.00023204621,0.00008021046,0.0020163544],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9962237,0.0014856567,0.00042083088,0.0004516484,0.0012180118,0.00020012993],"domain_scores_gemma":[0.99531734,0.002071273,0.00047599268,0.0012091692,0.0007594968,0.00016669363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045539257,0.0018299165,0.0008732699,0.003452426,0.002059496,0.005602787,0.0037206726,0.0029237585,0.0033877478],"category_scores_gemma":[0.010022305,0.0014073598,0.0031070374,0.0030789713,0.007938332,0.009951648,0.0035827206,0.00571843,0.0010781022],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009426807,0.00003623617,0.00019914258,0.00017334029,0.000027427832,0.00017130873,0.00092589593,0.016030118,0.0011024808,0.95838153,0.0015114244,0.021431787],"study_design_scores_gemma":[0.000017883225,0.000023205732,0.00007510445,0.00014316973,0.000033593336,0.0002083586,0.00026625616,0.04560067,0.0012334826,0.9179062,0.034466688,0.00002532558],"about_ca_topic_score_codex":0.005208308,"about_ca_topic_score_gemma":0.006477302,"teacher_disagreement_score":0.005602787,"about_ca_system_score_codex":0.0030669752,"about_ca_system_score_gemma":0.0044749156,"threshold_uncertainty_score":0.024083793},"labels":[],"label_agreement":null},{"id":"W2118944299","doi":"10.1145/581339.581390","title":"Concern graphs","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":289,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Source code; Java; Scalability; Usability; Software engineering; Abstraction; Representation (politics); Graph; Software; Programming language; Source lines of code; Theoretical computer science; Human–computer interaction; Database","score_opus":0.04323818430936585,"score_gpt":0.2629492601822788,"score_spread":0.21971107587291297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118944299","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01160749,0.0008435344,0.9207001,0.0013111943,0.0002331919,0.0006774925,0.008826366,0.01106708,0.044733644],"genre_scores_gemma":[0.19353415,0.00237657,0.6875424,0.0012272702,0.0002076762,0.0013297091,0.036637407,0.0047750864,0.07236979],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99816316,0.0004144685,0.00018563322,0.00041648193,0.0006515324,0.00016864519],"domain_scores_gemma":[0.99630475,0.0013770957,0.00034038114,0.000850621,0.0009467578,0.000180404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012609937,0.0011715189,0.0004943518,0.0033039073,0.0015515995,0.002391771,0.001572347,0.0013946221,0.016551415],"category_scores_gemma":[0.006493564,0.00068978406,0.001667495,0.0026783617,0.00085680437,0.005417434,0.0023076378,0.0014802509,0.0046930593],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019937758,0.0001960422,0.0057292064,0.00099678,0.00015124088,0.0012835958,0.0019688352,0.014855822,0.0072098263,0.55165964,0.09424629,0.32150328],"study_design_scores_gemma":[0.0000282188,0.00006582665,0.0015654567,0.00017777213,0.000106568164,0.0010929556,0.00048067048,0.025100652,0.0060690283,0.2986592,0.66659796,0.000055728688],"about_ca_topic_score_codex":0.0072076274,"about_ca_topic_score_gemma":0.009048946,"teacher_disagreement_score":0.016551415,"about_ca_system_score_codex":0.0010952934,"about_ca_system_score_gemma":0.0014647124,"threshold_uncertainty_score":0.055369973},"labels":[],"label_agreement":null},{"id":"W2118991740","doi":"10.1109/re.2005.17","title":"Concept identification in object-oriented domain analysis: why some students just don't get it","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Domain analysis; Object-oriented analysis and design; Domain (mathematical analysis); Identification (biology); Elevator; Task (project management); Object (grammar); Software engineering; Process (computing); Object-oriented programming; Software; Software system; Human–computer interaction; Programming language; Unified Modeling Language; Artificial intelligence; Engineering; Software construction; Systems engineering","score_opus":0.01553287274131132,"score_gpt":0.31044239566837295,"score_spread":0.29490952292706163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118991740","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63944674,0.0060976604,0.20148358,0.11982461,0.0009074486,0.0004014116,0.000089660185,0.00201466,0.029734282],"genre_scores_gemma":[0.89075977,0.001916195,0.07899765,0.016474763,0.00028347725,0.00029223767,0.00015508862,0.0006255778,0.010495233],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9626878,0.015408716,0.0018777011,0.004302142,0.013542282,0.002181378],"domain_scores_gemma":[0.8755077,0.06342802,0.008920196,0.012016791,0.03284619,0.0072810915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.055918876,0.00075545744,0.00080209196,0.00244917,0.004098008,0.010686299,0.002865586,0.0048221485,0.0031936518],"category_scores_gemma":[0.15359616,0.0012059353,0.0008874849,0.0021888013,0.010148847,0.020167245,0.007100566,0.008619707,0.003052797],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004889751,0.0010019636,0.10765733,0.0010175747,0.00010465671,0.0013014644,0.18277249,0.0023084083,0.008638286,0.0650321,0.049778298,0.5798985],"study_design_scores_gemma":[0.0003756494,0.0014789135,0.04607456,0.0023411617,0.0001679586,0.0063128737,0.23093632,0.021420319,0.014785384,0.381556,0.29392642,0.0006244744],"about_ca_topic_score_codex":0.0028362763,"about_ca_topic_score_gemma":0.0019383259,"teacher_disagreement_score":0.055918876,"about_ca_system_score_codex":0.0028760156,"about_ca_system_score_gemma":0.005062303,"threshold_uncertainty_score":0.29573077},"labels":[],"label_agreement":null},{"id":"W2119016375","doi":"10.1145/1985404.1985421","title":"Visualizing the evolution of code clones","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Visualization; Code (set theory); Software evolution; Set (abstract data type); Programming language; Scalability; Source code; Software; Software engineering; Theoretical computer science; Software development; Data mining; Database; Software construction; Biology; Genetics","score_opus":0.04860956927986116,"score_gpt":0.2898321808132298,"score_spread":0.24122261153336866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119016375","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5202353,0.002445193,0.43488047,0.0026044266,0.0002472264,0.00015415324,0.004709753,0.023400296,0.0113232],"genre_scores_gemma":[0.7765018,0.0010203919,0.2157071,0.0001638848,0.000079266894,0.00007386085,0.002333623,0.0011907596,0.002929236],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997038,0.00006390474,0.000022036353,0.00007488905,0.00010722462,0.00002799403],"domain_scores_gemma":[0.9966259,0.0017287432,0.00050188793,0.0003531822,0.0006141154,0.00017628647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007542001,0.00051784917,0.00029614198,0.0036589806,0.00045812989,0.0014861348,0.0004949987,0.00076304463,0.002300572],"category_scores_gemma":[0.0049093687,0.0002545119,0.00038931225,0.0019975244,0.00034278116,0.002121742,0.0011594633,0.0008705661,0.00030221153],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012029244,0.00021115375,0.05334672,0.0010765713,0.00021668017,0.0020971254,0.01998078,0.045634586,0.18164556,0.037217174,0.026872303,0.63049847],"study_design_scores_gemma":[0.0001758216,0.0005298384,0.09140127,0.0006245365,0.00041327515,0.004047369,0.0049393484,0.55766785,0.15807873,0.05713059,0.12464112,0.00035030933],"about_ca_topic_score_codex":0.0024486727,"about_ca_topic_score_gemma":0.002408607,"teacher_disagreement_score":0.0036589806,"about_ca_system_score_codex":0.0003639778,"about_ca_system_score_gemma":0.00037728852,"threshold_uncertainty_score":0.0076961517},"labels":[],"label_agreement":null},{"id":"W2119144699","doi":"10.1109/sbes.2009.28","title":"Software Artifact Metamodel","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artifact (error); Metamodeling; Software engineering; Software development; Goal-Driven Software Development Process; Unified Modeling Language; Software development process; Software; Software construction; Artificial intelligence; Programming language","score_opus":0.02222571836630289,"score_gpt":0.26905702773919954,"score_spread":0.24683130937289666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119144699","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004770755,0.00035453777,0.97345406,0.0005191615,0.00016583884,0.0006548781,0.0029606305,0.0044668796,0.012653291],"genre_scores_gemma":[0.06053865,0.0012372264,0.90075886,0.0005018989,0.00009156235,0.0015070102,0.012527103,0.001011004,0.021826621],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99737465,0.000739802,0.00058266596,0.00038228894,0.0007252128,0.00019530587],"domain_scores_gemma":[0.99667346,0.0007790878,0.00025846955,0.0012342727,0.0009073269,0.00014732349],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029199494,0.0010343889,0.00074839237,0.0033056291,0.00089927285,0.003395967,0.0020863872,0.0023768018,0.0047648097],"category_scores_gemma":[0.0042347433,0.0007101782,0.00197598,0.0024669925,0.00074145326,0.0028549915,0.0019890205,0.0025888903,0.00381759],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018718431,0.0003678623,0.0034894662,0.0013409667,0.00016872193,0.0014250387,0.0018684242,0.02961958,0.03457233,0.6634929,0.031284526,0.232183],"study_design_scores_gemma":[0.000048695223,0.00014035523,0.0008235532,0.0005421344,0.000101522666,0.001293333,0.0003447267,0.053701267,0.013645477,0.092536256,0.83674985,0.00007281658],"about_ca_topic_score_codex":0.003057026,"about_ca_topic_score_gemma":0.004088239,"teacher_disagreement_score":0.0047648097,"about_ca_system_score_codex":0.0013828961,"about_ca_system_score_gemma":0.003697614,"threshold_uncertainty_score":0.015939832},"labels":[],"label_agreement":null},{"id":"W2119281731","doi":"10.1145/2480362.2480574","title":"OSDC","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software engineering; Software development process; Software development; Process (computing); Software; Software quality; Software bug; Programming language","score_opus":0.012753415125360807,"score_gpt":0.23924327197206446,"score_spread":0.22648985684670367,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119281731","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015652817,0.0014307785,0.06615637,0.0023947828,0.002972223,0.00075393927,0.03870494,0.04082915,0.83110493],"genre_scores_gemma":[0.08311303,0.0018257883,0.04349089,0.0016641582,0.0007777763,0.0005522694,0.08794517,0.008916004,0.77171487],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986577,0.00015479699,0.00013532699,0.00028918314,0.0006072164,0.00015570964],"domain_scores_gemma":[0.9957742,0.0003443946,0.00020018182,0.0013307321,0.0018961838,0.00045434796],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001469147,0.0008909204,0.0006211714,0.0025260274,0.0011229808,0.0035832073,0.0013731178,0.0010265465,0.33607796],"category_scores_gemma":[0.004068988,0.00040093716,0.00047868222,0.0021091006,0.00044099937,0.0024657757,0.0024337603,0.0011685075,0.28080386],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045304684,0.00017947826,0.0037111787,0.00051615376,0.000026356509,0.0002926382,0.00034682913,0.0010945019,0.0061386684,0.033623513,0.4669219,0.48669568],"study_design_scores_gemma":[0.00003473603,0.00004990138,0.00091932586,0.00006472708,0.00000935705,0.00021784932,0.00006569147,0.0009709826,0.0021290404,0.002301112,0.99322295,0.000014373792],"about_ca_topic_score_codex":0.0027262464,"about_ca_topic_score_gemma":0.002448561,"teacher_disagreement_score":0.6639221,"about_ca_system_score_codex":0.0011941509,"about_ca_system_score_gemma":0.0020230196,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2119391892","doi":"10.1109/wcre.2003.1287245","title":"Detecting merging and splitting using origin analysis","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Abstraction; Code (set theory); Programming language; Plan (archaeology); Software; Software engineering; Static analysis; Software evolution; Theoretical computer science; Data mining; Software system; Software construction","score_opus":0.0290609720212598,"score_gpt":0.30160936359395846,"score_spread":0.27254839157269867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119391892","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33736885,0.00071951165,0.64350486,0.00027299192,0.00006944561,0.00020040688,0.00094146514,0.012666605,0.004255818],"genre_scores_gemma":[0.6120013,0.00028783717,0.38242757,0.00008492177,0.00004024717,0.00010503714,0.0018053184,0.0013223947,0.0019252836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950051,0.0007088287,0.0005649653,0.0011696522,0.0021625364,0.00038898102],"domain_scores_gemma":[0.97816133,0.007222303,0.0039098808,0.005273997,0.00497659,0.00045584777],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037436397,0.0007166481,0.0009267481,0.0075927763,0.0012563033,0.002708346,0.0022239243,0.0016691914,0.002192395],"category_scores_gemma":[0.018242523,0.00072144176,0.0010777881,0.004458699,0.0012991428,0.004967126,0.0044431183,0.0014141015,0.0007393207],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014544773,0.00032398934,0.25815585,0.00088431756,0.0002636646,0.005589413,0.012755467,0.014729871,0.09972421,0.041073788,0.0059886924,0.5590563],"study_design_scores_gemma":[0.00020270189,0.00060154515,0.086035095,0.0003212691,0.0008697248,0.0076931943,0.0045909453,0.43060318,0.31796676,0.059010625,0.09166761,0.00043728837],"about_ca_topic_score_codex":0.00274702,"about_ca_topic_score_gemma":0.0023798926,"teacher_disagreement_score":0.0075927763,"about_ca_system_score_codex":0.0008345271,"about_ca_system_score_gemma":0.0012594998,"threshold_uncertainty_score":0.019798458},"labels":[],"label_agreement":null},{"id":"W2119433210","doi":"10.1109/13.883346","title":"Experiences with a software maintenance project course","year":2000,"lang":"en","type":"article","venue":"IEEE Transactions on Education","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Course (navigation); Software engineering; Software maintenance; Software project management; Engineering management; Software; Software development; Engineering; Computer science; Software peer review; Systems engineering; Software construction; Programming language","score_opus":0.012746904188256481,"score_gpt":0.28094371734621243,"score_spread":0.26819681315795596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119433210","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9662335,0.0007153667,0.006336928,0.0046688714,0.00038784582,0.00028045784,0.00010410673,0.00035992858,0.020913059],"genre_scores_gemma":[0.92804635,0.0013212857,0.015886026,0.0024860105,0.00020384417,0.00023727638,0.00042504736,0.00029657938,0.051097486],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9959388,0.0015685753,0.00013365518,0.00037602804,0.0010498477,0.0009330418],"domain_scores_gemma":[0.98541033,0.0024807279,0.00034326166,0.00041863,0.0016562206,0.009690865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077651865,0.0010065329,0.0006324406,0.0011966638,0.009464813,0.0051466697,0.0021643846,0.0030479059,0.0073241405],"category_scores_gemma":[0.01304104,0.00076632085,0.00062278495,0.0011269039,0.0019003068,0.0033720774,0.006005292,0.0059409393,0.002287406],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082425796,0.027450752,0.040188313,0.00079770334,0.00008132543,0.012355186,0.42700312,0.0027345791,0.012466636,0.008818288,0.08401233,0.38326746],"study_design_scores_gemma":[0.00043409466,0.013213811,0.035625927,0.0006192571,0.00012369183,0.010125918,0.27692035,0.007197814,0.008758824,0.006343159,0.64027023,0.00036693865],"about_ca_topic_score_codex":0.0024334176,"about_ca_topic_score_gemma":0.0075847,"teacher_disagreement_score":0.009464813,"about_ca_system_score_codex":0.0018031304,"about_ca_system_score_gemma":0.0025556108,"threshold_uncertainty_score":0.041066706},"labels":[],"label_agreement":null},{"id":"W2119631871","doi":"10.1109/csmr.2007.9","title":"A Probabilistic Approach to Predict Changes in Object-Oriented Software Systems","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software maintenance; Software system; Software quality; Software metric; Software sizing; Software development; Reverse engineering; Software engineering; Source code; Unified Modeling Language; Reliability engineering; Software construction; Software; Programming language; Engineering","score_opus":0.019727696626792055,"score_gpt":0.25298655376447726,"score_spread":0.2332588571376852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119631871","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02510511,0.00014819202,0.9735082,0.000112713606,0.000009432948,0.000084717074,0.0001465541,0.0005343456,0.0003506392],"genre_scores_gemma":[0.48184896,0.00040851664,0.5156957,0.00006771648,0.000048977763,0.0003505145,0.0005796203,0.00008363591,0.000916418],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979861,0.00065316673,0.00015155364,0.000305541,0.00082864135,0.00007500003],"domain_scores_gemma":[0.9882841,0.008743199,0.0011964105,0.00062603335,0.0010450854,0.00010514088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026431899,0.00078657834,0.00059804943,0.003171831,0.0006290915,0.0009755749,0.0012325242,0.0010561238,0.0007948236],"category_scores_gemma":[0.01647016,0.0008284726,0.0010305704,0.0017324993,0.0006238641,0.0017046479,0.0007956158,0.0009613668,0.00025925884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011463184,0.00013696744,0.020390494,0.00019364365,0.00016347677,0.0002417461,0.00025142415,0.81889725,0.0048061074,0.010271482,0.00077786646,0.14375502],"study_design_scores_gemma":[0.000006456521,0.000033754117,0.0020698449,0.000009473588,0.000028458699,0.00010109393,0.000015291602,0.98897064,0.0008522402,0.0074082687,0.00048651308,0.000017951877],"about_ca_topic_score_codex":0.0066422583,"about_ca_topic_score_gemma":0.009013117,"teacher_disagreement_score":0.0066422583,"about_ca_system_score_codex":0.0007205661,"about_ca_system_score_gemma":0.0011578138,"threshold_uncertainty_score":0.01397866},"labels":[],"label_agreement":null},{"id":"W2119686964","doi":"10.1109/wcre.2008.12","title":"Detecting Clones in Business Applications","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Business rule; Business process modeling; Computer science; Artifact-centric business process model; Business process; Business process discovery; Business process management; Business domain; Business Process Model and Notation; clone (Java method); Source code; Business logic; Software engineering; Open source; Business analysis; Database; Business model; Software; Operating system; Business; Work in process; Marketing","score_opus":0.025802364406049178,"score_gpt":0.2617213747519449,"score_spread":0.2359190103458957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119686964","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.71965617,0.0013773644,0.26717502,0.00043121626,0.000075223994,0.00058946013,0.00046548512,0.007167919,0.0030620901],"genre_scores_gemma":[0.7464071,0.00060224894,0.24785133,0.0002516714,0.00004108863,0.00028160616,0.001202064,0.0007231748,0.0026397244],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99494815,0.00088511285,0.00042647123,0.0011002986,0.0023553201,0.00028462586],"domain_scores_gemma":[0.9567356,0.022723056,0.006727114,0.0043971734,0.008955058,0.0004621042],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022948366,0.00076148764,0.00079393183,0.0055053905,0.0014389388,0.0020990488,0.0013947776,0.001972179,0.00053873064],"category_scores_gemma":[0.03816477,0.00085895386,0.00080994534,0.003263296,0.00094877725,0.0032871142,0.0018071445,0.0009787467,0.00036701807],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040455302,0.00048046804,0.24442005,0.0009252807,0.00021395415,0.0060933563,0.011381489,0.011401217,0.12426343,0.0074990676,0.004563996,0.5883533],"study_design_scores_gemma":[0.00015078078,0.0010599969,0.20627311,0.0005241324,0.0010967061,0.015297387,0.0057445266,0.396218,0.3007214,0.025303097,0.047246564,0.000364264],"about_ca_topic_score_codex":0.004881008,"about_ca_topic_score_gemma":0.0051908772,"teacher_disagreement_score":0.0055053905,"about_ca_system_score_codex":0.0010710873,"about_ca_system_score_gemma":0.0013772912,"threshold_uncertainty_score":0.0121364},"labels":[],"label_agreement":null},{"id":"W2119731669","doi":"10.1109/msr.2007.25","title":"Predicting Eclipse Bug Lifetimes","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":176,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Eclipse; Computer science; Software bug; Process (computing); Task (project management); Data mining; Software; Data modeling; Software engineering; Programming language; Engineering; Systems engineering","score_opus":0.013506405826594937,"score_gpt":0.27023283710525287,"score_spread":0.2567264312786579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119731669","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9654111,0.00047761982,0.02670338,0.00017190461,0.000024469682,0.000047137946,0.004661189,0.0017574048,0.0007456202],"genre_scores_gemma":[0.9724336,0.00022103403,0.019069834,0.000017446544,0.000009401596,0.00004987232,0.00755227,0.00007148617,0.00057509815],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99934393,0.00012312763,0.00007202246,0.00019976594,0.00020235138,0.000058691716],"domain_scores_gemma":[0.98769766,0.0075090206,0.0022800486,0.0008082363,0.0014204368,0.0002845015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014322127,0.00082720653,0.0004561213,0.0031544636,0.00022339867,0.0006221994,0.0005097191,0.00066183554,0.00050909945],"category_scores_gemma":[0.0147125,0.00034501185,0.00049315044,0.0014632117,0.00015093035,0.0009873741,0.00036417827,0.0006643131,0.00031702724],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061184546,0.00045008378,0.49495137,0.00032682755,0.0001441801,0.00035249707,0.00038130675,0.29012328,0.0067552854,0.0011299221,0.0049791117,0.1997943],"study_design_scores_gemma":[0.000028733119,0.00033696086,0.08847749,0.000035492845,0.00005383936,0.00021296887,0.00015174443,0.90238655,0.0051173987,0.0014564028,0.0017143002,0.000028091457],"about_ca_topic_score_codex":0.008729492,"about_ca_topic_score_gemma":0.011798085,"teacher_disagreement_score":0.008729492,"about_ca_system_score_codex":0.0007617797,"about_ca_system_score_gemma":0.0005669179,"threshold_uncertainty_score":0.01735735},"labels":[],"label_agreement":null},{"id":"W2119836629","doi":"10.1109/wcre.2000.891469","title":"Applying traditional Unix tools during maintenance: an experience report","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Unix; Computer science; Documentation; Software engineering; Software maintenance; Unix architecture; Set (abstract data type); Software; Operating system; Software development; Programming language","score_opus":0.08522765039597292,"score_gpt":0.2772998073321017,"score_spread":0.19207215693612878,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119836629","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94759977,0.002900389,0.032516804,0.00095669035,0.00006660866,0.0003645877,0.00012264978,0.0010575155,0.014414898],"genre_scores_gemma":[0.8955237,0.0047952808,0.07633447,0.00050475803,0.00010571114,0.0002248415,0.00061412784,0.0005530217,0.021344101],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99219584,0.0033499654,0.0006353793,0.0008089589,0.0022213955,0.0007884569],"domain_scores_gemma":[0.9856339,0.007193079,0.0008154508,0.002242725,0.0024826054,0.0016322943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011736577,0.0009975947,0.00063494546,0.0014708596,0.0023033768,0.0024435765,0.003204312,0.0024109902,0.0024533023],"category_scores_gemma":[0.024793098,0.0007986074,0.0005065838,0.0014864497,0.0018908891,0.0041642804,0.0024526126,0.0019867187,0.0014678828],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061601953,0.008402086,0.022033077,0.00060416054,0.000069137626,0.0020344423,0.09889211,0.0012199002,0.02575803,0.0015310815,0.007943488,0.83089644],"study_design_scores_gemma":[0.001349617,0.051858664,0.22017895,0.0012332622,0.000857241,0.03526542,0.097842045,0.029977435,0.17647809,0.006181912,0.37792283,0.00085448555],"about_ca_topic_score_codex":0.0037584295,"about_ca_topic_score_gemma":0.005649495,"teacher_disagreement_score":0.011736577,"about_ca_system_score_codex":0.000794009,"about_ca_system_score_gemma":0.0013321039,"threshold_uncertainty_score":0.062069654},"labels":[],"label_agreement":null},{"id":"W2120167743","doi":"10.1109/icsm.2007.4362634","title":"Polylingual Dependency Analysis Using Island Grammars: A Cost Versus Accuracy Evaluation","year":2007,"lang":"en","type":"article","venue":"Proceedings/Proceedings - Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Dependency (UML); Rule-based machine translation; Computer science; Natural language processing; L-attributed grammar; Artificial intelligence; Programming language; Context-free grammar","score_opus":0.07205561844229989,"score_gpt":0.3479809792130327,"score_spread":0.2759253607707328,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120167743","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57749736,0.0002859672,0.40716693,0.0003012164,0.000049233655,0.00046945518,0.0007124821,0.008542768,0.0049745645],"genre_scores_gemma":[0.60245275,0.00016338032,0.3927379,0.000046105342,0.000016443331,0.00014905045,0.0011652366,0.001480099,0.0017890048],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948061,0.0017808896,0.00048961374,0.00063173013,0.0020792037,0.00021239756],"domain_scores_gemma":[0.94568235,0.038812824,0.0013229136,0.008002087,0.00590294,0.0002768953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046974756,0.00061926205,0.00069052074,0.0020516377,0.0005731898,0.0015044886,0.0017846182,0.00070408854,0.0022779526],"category_scores_gemma":[0.03260977,0.00034476115,0.00070863875,0.001843957,0.00074530486,0.0028544974,0.0015395384,0.0008414308,0.0005967873],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017226984,0.00057362305,0.038989715,0.0006995204,0.00023053923,0.0008508288,0.0036174937,0.14396307,0.058304504,0.01485905,0.0034796013,0.7327092],"study_design_scores_gemma":[0.000118005744,0.00065044285,0.01511935,0.000049997376,0.00018805005,0.00083787966,0.0009998922,0.8571236,0.10542699,0.012420859,0.006925223,0.00013970908],"about_ca_topic_score_codex":0.0046325754,"about_ca_topic_score_gemma":0.0062690172,"teacher_disagreement_score":0.0046974756,"about_ca_system_score_codex":0.0010295303,"about_ca_system_score_gemma":0.0009155052,"threshold_uncertainty_score":0.024842918},"labels":[],"label_agreement":null},{"id":"W2120202833","doi":"10.5555/2337223.2337485","title":"On the analysis of evolution of software artefacts and programs","year":2012,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Macro; Computer science; Software evolution; Construct (python library); Field (mathematics); Tree (set theory); Software; Data science; Artificial intelligence; Theoretical computer science; Software development; Programming language; Mathematics","score_opus":0.03546389148086995,"score_gpt":0.2743082206007366,"score_spread":0.23884432911986667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120202833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045625925,0.0069445195,0.94025016,0.0006649613,0.000056969464,0.0001673385,0.00064405927,0.0011140383,0.004532154],"genre_scores_gemma":[0.22578683,0.005328209,0.7630875,0.00018840862,0.000117437165,0.0003058343,0.001742647,0.0004010168,0.0030421605],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99561286,0.0015033403,0.00033829737,0.00093181763,0.0014037598,0.0002100336],"domain_scores_gemma":[0.98447627,0.010509238,0.00204422,0.0015009271,0.0012426757,0.00022662555],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029666557,0.0011795716,0.0009909256,0.011981526,0.001170384,0.0034482484,0.0012549686,0.0012705697,0.0015629355],"category_scores_gemma":[0.01751466,0.0005834967,0.0018086634,0.009607242,0.0034551374,0.003875369,0.0017014844,0.0013425408,0.00047345995],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022638425,0.00019297047,0.051761333,0.0023861944,0.00036296254,0.0013987683,0.004758496,0.10290937,0.019492175,0.2459311,0.0024039452,0.5681764],"study_design_scores_gemma":[0.0000364653,0.00025583233,0.055717956,0.0011922688,0.0002711863,0.0022258481,0.0016971626,0.4163063,0.019632945,0.42418697,0.07825515,0.00022196698],"about_ca_topic_score_codex":0.0055200295,"about_ca_topic_score_gemma":0.0030971577,"teacher_disagreement_score":0.011981526,"about_ca_system_score_codex":0.0018728675,"about_ca_system_score_gemma":0.0014412167,"threshold_uncertainty_score":0.015689373},"labels":[],"label_agreement":null},{"id":"W2120293505","doi":"10.1109/iceccs.2011.36","title":"Analyzing and Forecasting Near-Miss Clones in Evolving Software: An Empirical Study","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Software evolution; Computer science; clone (Java method); Software maintenance; Programming language; Java; Software system; Dependency (UML); Software development; Software; Code (set theory); Software engineering; Software construction; Biology","score_opus":0.08643892818750509,"score_gpt":0.31876768262386607,"score_spread":0.23232875443636097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120293505","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99911505,0.00004809356,0.00067761407,0.000014443722,8.704409e-7,0.0000074148666,0.000047275433,0.000008423293,0.000080837664],"genre_scores_gemma":[0.9983485,0.000059845537,0.0011986156,0.000007668319,0.000003129235,0.000012269258,0.00025457688,0.0000059583317,0.00010957226],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9971118,0.0011613957,0.0002852992,0.00053745246,0.0007618906,0.00014215196],"domain_scores_gemma":[0.8750874,0.09624856,0.0152991265,0.0062881033,0.005840169,0.0012366114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049246736,0.00029599323,0.00036471034,0.0022651174,0.00042352447,0.0008199398,0.00086809014,0.0009630182,0.00045172553],"category_scores_gemma":[0.04428581,0.00030439094,0.0004383188,0.002037052,0.00089656806,0.0017322288,0.00058055937,0.0011332446,0.00021577242],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016190755,0.00036617776,0.9649962,0.00004060858,0.00007621148,0.00032325645,0.0008435384,0.011186876,0.001138412,0.00024224485,0.00017687182,0.0204476],"study_design_scores_gemma":[0.000018083581,0.00068251765,0.84443057,0.000023976301,0.00006521202,0.0011308638,0.0012530653,0.14865051,0.0024383937,0.0005731431,0.00069330144,0.000040298448],"about_ca_topic_score_codex":0.0035425727,"about_ca_topic_score_gemma":0.0034200198,"teacher_disagreement_score":0.0049246736,"about_ca_system_score_codex":0.0005560438,"about_ca_system_score_gemma":0.0003100596,"threshold_uncertainty_score":0.026044428},"labels":[],"label_agreement":null},{"id":"W2120459757","doi":"10.1109/csmr.2010.19","title":"Utilizing Debug Information to Compact Loops in Large Program Traces","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; University of Victoria","funders":"","keywords":"Computer science; Debugging; Java; Programming language; Sequence diagram; Source code; Software visualization; Software; Program comprehension; Unified Modeling Language; Context (archaeology); Visualization; Software system; Software engineering; Software construction; Data mining","score_opus":0.015297985088313494,"score_gpt":0.30475029246576846,"score_spread":0.289452307377455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120459757","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054256324,0.00019033172,0.93642724,0.00014954376,0.000026084263,0.00008821929,0.00013345687,0.007896911,0.00083180977],"genre_scores_gemma":[0.20504475,0.00022508268,0.7925501,0.000027484253,0.000014094177,0.00007672432,0.00038797787,0.00082211505,0.00085164094],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993794,0.00015408405,0.000059386286,0.00012729858,0.00023095025,0.00004891495],"domain_scores_gemma":[0.99273926,0.0047397325,0.0006962927,0.0008879159,0.0007442441,0.0001924875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013392721,0.0012024203,0.00073765416,0.003933279,0.0005490059,0.0017302189,0.0009093445,0.0006197329,0.002293711],"category_scores_gemma":[0.011597067,0.0005264435,0.0003751193,0.0019656967,0.0006431518,0.0028107907,0.0018096432,0.0007852924,0.00044823484],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073902425,0.00018198471,0.0067720874,0.00040512648,0.00005962751,0.00045667117,0.002606608,0.041740462,0.0626934,0.02053433,0.0022707717,0.8615399],"study_design_scores_gemma":[0.0001270458,0.0004996596,0.005463255,0.00021430936,0.000099714154,0.0012845974,0.0010624517,0.73681366,0.17273533,0.051059432,0.030485038,0.00015554536],"about_ca_topic_score_codex":0.0013855036,"about_ca_topic_score_gemma":0.0018801814,"teacher_disagreement_score":0.003933279,"about_ca_system_score_codex":0.00044652112,"about_ca_system_score_gemma":0.000747136,"threshold_uncertainty_score":0.007673204},"labels":[],"label_agreement":null},{"id":"W2120738100","doi":"10.1109/tse.2006.102","title":"Empirical Analysis of Object-Oriented Design Metrics for Predicting High and Low Severity Faults","year":2006,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":349,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Larner College of Medicine, University of Vermont; University of Alberta","keywords":"Computer science; Fault (geology); Data mining; Software fault tolerance; Empirical research; Object-oriented programming; Machine learning; Reliability engineering; Logistic regression; Artificial intelligence; Fault tolerance; Distributed computing; Engineering; Statistics; Mathematics; Programming language","score_opus":0.018279048185542823,"score_gpt":0.260653268209901,"score_spread":0.24237422002435816,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120738100","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9877809,0.0003834303,0.010417827,0.00013592624,0.000007811094,0.000026371798,0.0006326763,0.0001194547,0.00049562607],"genre_scores_gemma":[0.99547064,0.00008215173,0.003403874,0.000011003432,0.0000075250014,0.000017168843,0.00092830666,0.000013093194,0.00006630466],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9964851,0.0016256283,0.00033331005,0.00033119216,0.0010808414,0.00014388216],"domain_scores_gemma":[0.89432627,0.08083223,0.013834485,0.0042202175,0.005693038,0.0010937087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068428507,0.0008595249,0.00041916792,0.004498734,0.00020297998,0.0006259704,0.00043189412,0.0006175622,0.0004445636],"category_scores_gemma":[0.060883675,0.00016787206,0.00039480889,0.0029093192,0.00044463226,0.0011643238,0.00046845397,0.0006937869,0.00024144241],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012037052,0.00014060948,0.94749177,0.000052971845,0.00013375249,0.000053105276,0.000101357306,0.022021312,0.00063269626,0.00020912803,0.0004159277,0.028627088],"study_design_scores_gemma":[0.000035408026,0.0004128474,0.71041214,0.000038068138,0.00007028013,0.00025630955,0.00019231495,0.2830052,0.0027668553,0.0018722058,0.00090434076,0.000034007106],"about_ca_topic_score_codex":0.0019360883,"about_ca_topic_score_gemma":0.0020999638,"teacher_disagreement_score":0.0068428507,"about_ca_system_score_codex":0.000403286,"about_ca_system_score_gemma":0.00039549242,"threshold_uncertainty_score":0.0361889},"labels":[],"label_agreement":null},{"id":"W2120755390","doi":"10.1109/tse.2014.2387172","title":"Extracting Development Tasks to Navigate Software Documentation","year":2014,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Software documentation; Computer science; Internal documentation; Software engineering; Software development; Application programming interface; Technical documentation; Software; Task (project management); World Wide Web; Software construction; Programming language; Systems engineering; Engineering","score_opus":0.013513800978778162,"score_gpt":0.25689729459795757,"score_spread":0.2433834936191794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120755390","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40441975,0.0028490087,0.53192824,0.0017740333,0.0002091743,0.0026305034,0.017481927,0.025296109,0.013411326],"genre_scores_gemma":[0.253318,0.0012263015,0.714933,0.00025362446,0.000057132314,0.0011740024,0.022569412,0.0019078271,0.0045607863],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967,0.0010436119,0.00064453087,0.0004638204,0.0009253863,0.00022267828],"domain_scores_gemma":[0.9675826,0.020962434,0.0028368304,0.0024085983,0.005583963,0.00062560197],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035798384,0.0017197339,0.0009992875,0.011277053,0.0012413196,0.002248227,0.0011382418,0.0011491199,0.0023018003],"category_scores_gemma":[0.037713587,0.00077417534,0.0009044351,0.0050535686,0.00037208622,0.0031590192,0.0021989872,0.0010306053,0.002167296],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064611866,0.0005117638,0.0376596,0.005005767,0.000099633005,0.0012736638,0.015810113,0.004753964,0.033847585,0.00577627,0.048684627,0.845931],"study_design_scores_gemma":[0.0005365692,0.0016878757,0.102924235,0.00500647,0.0005136765,0.0043935087,0.03144958,0.22040917,0.13178794,0.03643457,0.46407583,0.00078055327],"about_ca_topic_score_codex":0.0059148134,"about_ca_topic_score_gemma":0.012246673,"teacher_disagreement_score":0.011277053,"about_ca_system_score_codex":0.00095525576,"about_ca_system_score_gemma":0.0037601423,"threshold_uncertainty_score":0.018932223},"labels":[],"label_agreement":null},{"id":"W2120801660","doi":"10.1049/iet-sen.2012.0058","title":"Conflict‐aware optimal scheduling of prioritised code clone refactoring","year":2013,"lang":"en","type":"article","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Maintainability; Programming language; Source code; Scheduling (production processes); Software maintenance; Software engineering; Software; Software system; Engineering","score_opus":0.026973757101524092,"score_gpt":0.2775546649302163,"score_spread":0.2505809078286922,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120801660","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17693704,0.00050773233,0.816798,0.0003089099,0.00007852326,0.00022251895,0.00017991234,0.00070330064,0.004264119],"genre_scores_gemma":[0.7886105,0.00017459015,0.20938186,0.000081688726,0.000020583178,0.0001260358,0.00020472618,0.00013174534,0.0012682973],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99888545,0.00038776852,0.00005669763,0.00021207717,0.00026596527,0.00019197582],"domain_scores_gemma":[0.9969348,0.0016905392,0.0004068895,0.00021913587,0.00049684034,0.0002517845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016051739,0.00085610006,0.0008815163,0.00082584715,0.00042371184,0.00083097094,0.00142174,0.0006071383,0.001498269],"category_scores_gemma":[0.0048342342,0.00059345557,0.00053480675,0.00088368513,0.0004357857,0.00071451,0.000619062,0.0007820972,0.00021614801],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002689643,0.00012393914,0.0012647521,0.00009675855,0.000050592203,0.00009499953,0.000106654705,0.9283394,0.0057432833,0.0043236474,0.00086880714,0.058718164],"study_design_scores_gemma":[0.00002332211,0.00006344325,0.0003673689,0.000004780754,0.000013238512,0.000015928843,0.000019584928,0.99640876,0.0012366482,0.0014309802,0.0004093992,0.000006586257],"about_ca_topic_score_codex":0.011366065,"about_ca_topic_score_gemma":0.009859888,"teacher_disagreement_score":0.011366065,"about_ca_system_score_codex":0.0013451965,"about_ca_system_score_gemma":0.0026953367,"threshold_uncertainty_score":0.022599816},"labels":[],"label_agreement":null},{"id":"W2120819709","doi":"10.1109/wcre.2004.20","title":"Exploring software evolution using spectrographs","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Visualization; Software visualization; Software; Spectrograph; Software development; Scalability; Software maintenance; Software analytics; Software system; Software engineering; Component-based software engineering; Artificial intelligence; Software construction; Programming language; Astronomy; Database","score_opus":0.09744252614984274,"score_gpt":0.280585830512499,"score_spread":0.1831433043626563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120819709","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6693816,0.00072930666,0.31255692,0.0006892791,0.000026923168,0.00007165427,0.00068337197,0.005114538,0.010746451],"genre_scores_gemma":[0.8642292,0.000358432,0.13329987,0.000053709948,0.00001633655,0.00003952912,0.0003651891,0.00032784184,0.0013098227],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997787,0.00008072036,0.000008795898,0.000036458234,0.000073978,0.000021301235],"domain_scores_gemma":[0.9987143,0.0008025395,0.00012512354,0.00012288119,0.0001699463,0.00006525885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005915169,0.0003129534,0.00017668246,0.0027536522,0.00044046264,0.00086563244,0.00031186003,0.0005122972,0.0019015182],"category_scores_gemma":[0.0022915516,0.00020528877,0.00022548041,0.0013541592,0.0003279718,0.0012374064,0.0009052594,0.0005402499,0.00018779653],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005094076,0.00018435564,0.041085098,0.00048238484,0.00012794981,0.0017212834,0.01560663,0.04926468,0.22820924,0.025934296,0.009882472,0.6269922],"study_design_scores_gemma":[0.00012178607,0.00032633162,0.1143437,0.00020809394,0.0001681871,0.0026767447,0.005772067,0.6452442,0.10823821,0.047459822,0.07520359,0.00023728478],"about_ca_topic_score_codex":0.0028204557,"about_ca_topic_score_gemma":0.0029859475,"teacher_disagreement_score":0.0028204557,"about_ca_system_score_codex":0.00043946112,"about_ca_system_score_gemma":0.00023782522,"threshold_uncertainty_score":0.0063611865},"labels":[],"label_agreement":null},{"id":"W2121042490","doi":"","title":"On specifying systems that connect to the physical world","year":2006,"lang":"en","type":"article","venue":"New Trends in Software Methodologies, Tools and Techniques","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Domain (mathematical analysis); Base (topology); Software; Software engineering; Industrial engineering; Risk analysis (engineering); Engineering; Mathematics; Programming language","score_opus":0.17654207095559094,"score_gpt":0.37930087428864767,"score_spread":0.20275880333305674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121042490","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035217914,0.0038097803,0.9232046,0.0053210864,0.0006260696,0.00024145606,0.00021457695,0.00075347663,0.06230715],"genre_scores_gemma":[0.111948624,0.018202754,0.8273376,0.0043166648,0.001533428,0.0016181574,0.0011036717,0.0009779383,0.032961212],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99264103,0.0032891459,0.0005888282,0.00078460766,0.0022996222,0.00039673343],"domain_scores_gemma":[0.9909758,0.005296985,0.00060345005,0.0021199877,0.0007052743,0.00029852416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057917503,0.0032886348,0.001350039,0.0022881068,0.0032821158,0.007208474,0.0034705636,0.0063581457,0.010765853],"category_scores_gemma":[0.014197437,0.0016367782,0.002009418,0.0037031532,0.023093028,0.019677509,0.0083905645,0.0075728353,0.004052332],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012871397,0.000009344256,0.00009232209,0.00021429386,0.000008369063,0.00014648051,0.00089911226,0.0019789254,0.00048314026,0.98228776,0.0018631186,0.012004243],"study_design_scores_gemma":[0.000030527754,0.000055414377,0.00020266375,0.0004791242,0.000029217732,0.0003965671,0.00050907675,0.005810502,0.0014185276,0.8371322,0.1538825,0.00005366931],"about_ca_topic_score_codex":0.00461609,"about_ca_topic_score_gemma":0.0057779886,"teacher_disagreement_score":0.010765853,"about_ca_system_score_codex":0.0020851307,"about_ca_system_score_gemma":0.0038073165,"threshold_uncertainty_score":0.03601533},"labels":[],"label_agreement":null},{"id":"W2121111390","doi":"10.1109/ms.2008.51","title":"The Software Engineering Silver Bullet Conundrum","year":2008,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Silver bullet; Software; Computer science; Software engineering; Social software engineering; Engineering; Software development; Software construction; Materials science; Programming language","score_opus":0.01639189029657439,"score_gpt":0.22654515245338033,"score_spread":0.21015326215680594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121111390","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051488252,0.030753672,0.0664576,0.7573046,0.0113853365,0.0000547955,0.00006500998,0.00066699256,0.1281632],"genre_scores_gemma":[0.24928117,0.034575522,0.08428429,0.4691315,0.01718758,0.00052579696,0.0001715046,0.0019054262,0.14293715],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9709952,0.010493734,0.0012737382,0.0032361208,0.012485665,0.0015155276],"domain_scores_gemma":[0.9168041,0.05717811,0.0024147597,0.006545476,0.012687442,0.004370039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028840132,0.0012928909,0.0012971831,0.0035712027,0.0065454575,0.012808034,0.0026221268,0.012790079,0.009378058],"category_scores_gemma":[0.08318667,0.00094318064,0.0010348625,0.0018672895,0.03673144,0.030623745,0.011866835,0.02615483,0.005631969],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002702339,0.00003283881,0.0002228805,0.000083253064,0.000009393512,0.00007920665,0.0008836507,0.00025705821,0.00013627335,0.86332643,0.09846125,0.03648077],"study_design_scores_gemma":[0.000024809213,0.000038445276,0.0001094231,0.00033479524,0.0000062522604,0.00015951759,0.0005784597,0.00070180185,0.00030476667,0.6700304,0.32768148,0.000029847843],"about_ca_topic_score_codex":0.00328633,"about_ca_topic_score_gemma":0.00232819,"teacher_disagreement_score":0.028840132,"about_ca_system_score_codex":0.0056415875,"about_ca_system_score_gemma":0.008679628,"threshold_uncertainty_score":0.15252304},"labels":[],"label_agreement":null},{"id":"W2121678797","doi":"10.1109/ast.1996.506477","title":"Using software metrics tools for maintenance decisions: a classroom exercise","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Class (philosophy); Software engineering; Software maintenance; Software; Software quality; Directory; Component (thermodynamics); Quality (philosophy); Software system; Software development; Programming language; Artificial intelligence; Operating system","score_opus":0.13562495081796458,"score_gpt":0.3181187319792758,"score_spread":0.18249378116131124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121678797","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7468992,0.00076124765,0.2078114,0.012005044,0.00021569061,0.0010576565,0.00016262483,0.002692099,0.028395064],"genre_scores_gemma":[0.5888734,0.0009763574,0.39487845,0.0012430481,0.00012772564,0.00059193425,0.000216468,0.0003303527,0.012762289],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962968,0.002098203,0.00018497187,0.0004418936,0.0006595895,0.0003185575],"domain_scores_gemma":[0.9796165,0.015838806,0.00074956944,0.0012728366,0.001461848,0.0010604003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073614093,0.0016156551,0.0010308945,0.0012105855,0.0025997013,0.002845861,0.0022994997,0.0026595981,0.002625703],"category_scores_gemma":[0.02085533,0.00072852825,0.00044976195,0.00093796675,0.001458684,0.0047281533,0.0025525838,0.00340417,0.00130492],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058829924,0.020698862,0.012905908,0.00086232595,0.00006853351,0.0019805494,0.08838095,0.008290504,0.03444742,0.021290375,0.033513743,0.77697253],"study_design_scores_gemma":[0.001737214,0.017232634,0.0407494,0.0023161327,0.00028809294,0.008134843,0.10879289,0.13529688,0.11695524,0.16027705,0.4075104,0.00070926937],"about_ca_topic_score_codex":0.0009276258,"about_ca_topic_score_gemma":0.0025925806,"teacher_disagreement_score":0.0073614093,"about_ca_system_score_codex":0.001347352,"about_ca_system_score_gemma":0.0014542795,"threshold_uncertainty_score":0.03893131},"labels":[],"label_agreement":null},{"id":"W2121820828","doi":"10.1109/csmr.2006.3","title":"A framework for software architecture refactoring using model transformations and semantic annotations","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code refactoring; Computer science; Maintainability; Software engineering; Unified Modeling Language; Software architecture; Software system; Software architecture description; Software evolution; Programming language; Context (archaeology); Reference architecture; Software; Software construction","score_opus":0.03295478639553996,"score_gpt":0.298517837636581,"score_spread":0.26556305124104107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121820828","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00026647982,0.000076696284,0.99702734,0.00012655443,0.000017751934,0.00011427005,0.000051257506,0.0016335556,0.0006861442],"genre_scores_gemma":[0.0062301164,0.00018286234,0.9922609,0.000044069024,0.0000211226,0.00026810335,0.0002892763,0.00020668856,0.0004968735],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9925288,0.0029091428,0.0010781325,0.00084169564,0.0023171469,0.00032510937],"domain_scores_gemma":[0.9930692,0.0026539243,0.0008095394,0.0021075243,0.0010894343,0.00027046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012427794,0.0026216456,0.0016940903,0.0054151276,0.0022040682,0.0060851336,0.005878092,0.0036201268,0.0035491008],"category_scores_gemma":[0.014174591,0.002164463,0.0053365347,0.0037738748,0.004333086,0.006063413,0.0048375637,0.0052532693,0.0019406482],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001172186,0.00026769447,0.00091591623,0.0009471248,0.00026975697,0.0013238423,0.002592708,0.099788584,0.007502234,0.62217546,0.00790488,0.25619465],"study_design_scores_gemma":[0.00012362756,0.0001730273,0.00047009147,0.0008752135,0.00023800593,0.00096124224,0.0004821036,0.38269955,0.0085741645,0.45996758,0.14521872,0.00021668735],"about_ca_topic_score_codex":0.013522349,"about_ca_topic_score_gemma":0.016477307,"teacher_disagreement_score":0.013522349,"about_ca_system_score_codex":0.0024856403,"about_ca_system_score_gemma":0.005861708,"threshold_uncertainty_score":0.06572521},"labels":[],"label_agreement":null},{"id":"W2121856880","doi":"10.1177/1063293x0100900403","title":"Project Estimation from Feasibility Study until Completion: A Quantitative Methodology","year":2001,"lang":"en","type":"article","venue":"Concurrent Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Concordia University","funders":"","keywords":"Estimation; Computer science; Project plan; Project management; Plan (archaeology); Project planning; Project management triangle; Systems engineering; Operations research; Risk analysis (engineering); Industrial engineering; Process management; Engineering; Business","score_opus":0.2436359053915479,"score_gpt":0.41609952990918836,"score_spread":0.17246362451764047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121856880","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042097326,0.00023809343,0.9925364,0.0001632676,0.000015969255,0.00040871082,0.00017380506,0.00022849493,0.002025499],"genre_scores_gemma":[0.12176837,0.0005397351,0.87319773,0.00008087549,0.000040230014,0.0025835016,0.00060421217,0.00016254278,0.0010227127],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9605896,0.02602286,0.0026355593,0.0022887024,0.007979653,0.00048363672],"domain_scores_gemma":[0.82402366,0.12978667,0.015036231,0.0085632205,0.021478921,0.001111203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.043489672,0.0017026673,0.00095827715,0.0060635945,0.0009872053,0.003411848,0.002060845,0.0010912883,0.004034786],"category_scores_gemma":[0.15293013,0.0013371202,0.00095181895,0.0045020347,0.001888378,0.0076278374,0.0025064312,0.0016998678,0.0009013888],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037118405,0.0003763328,0.011948244,0.0023843173,0.0002911743,0.00018802263,0.0029995157,0.13900837,0.0072436314,0.122063056,0.0047676363,0.7083585],"study_design_scores_gemma":[0.00035143283,0.005519003,0.026027732,0.0028900374,0.00065536995,0.0011915119,0.0045521725,0.51247495,0.042234406,0.29800668,0.10535623,0.0007403976],"about_ca_topic_score_codex":0.0017811193,"about_ca_topic_score_gemma":0.0014768983,"teacher_disagreement_score":0.043489672,"about_ca_system_score_codex":0.0014607871,"about_ca_system_score_gemma":0.005284128,"threshold_uncertainty_score":0.22999817},"labels":[],"label_agreement":null},{"id":"W2121883265","doi":"10.1109/ithet.2005.1560247","title":"Experience in Teaching an Introductory Software Engineering Course","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Acadia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Personal software process; Software engineering; Software development; Software Engineering Process Group; Social software engineering; Computer science; Software peer review; Curriculum; Team software process; Engineering management; Software project management; Software development process; Process (computing); Course (navigation); Work (physics); Software; Quality (philosophy); Software construction; Engineering; Pedagogy; Programming language","score_opus":0.016417555130620032,"score_gpt":0.2911493147057234,"score_spread":0.2747317595751034,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121883265","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9560772,0.0016312231,0.009797518,0.0039183307,0.00062137196,0.00051879656,0.00027703075,0.00076001574,0.026398532],"genre_scores_gemma":[0.9236611,0.0019708911,0.016200919,0.0060813352,0.00030887837,0.00024238406,0.0006408696,0.00028799198,0.050605662],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971,0.0006874315,0.00014641204,0.0004394668,0.0006641518,0.0009624365],"domain_scores_gemma":[0.9846903,0.0018078461,0.00031408167,0.0003313247,0.0021890278,0.010667407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032029217,0.001274852,0.0012074986,0.0015312093,0.005938361,0.0046550324,0.0015552967,0.002900367,0.02164812],"category_scores_gemma":[0.009357353,0.0009829093,0.0012050465,0.0014792369,0.0014914747,0.0026564393,0.00332176,0.003342613,0.006041351],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018534281,0.047005724,0.08316634,0.0018143132,0.00016692198,0.034355555,0.27563357,0.0034284052,0.047413852,0.004581046,0.06450848,0.4360725],"study_design_scores_gemma":[0.0008706672,0.03312048,0.127639,0.0010323251,0.00039460135,0.047323722,0.19983868,0.010020072,0.024602836,0.0052404744,0.5492107,0.00070635777],"about_ca_topic_score_codex":0.0028769635,"about_ca_topic_score_gemma":0.0071452754,"teacher_disagreement_score":0.02164812,"about_ca_system_score_codex":0.0019878044,"about_ca_system_score_gemma":0.003117377,"threshold_uncertainty_score":0.07242012},"labels":[],"label_agreement":null},{"id":"W2122010595","doi":"10.1109/icsm.2005.13","title":"A reference architecture for Web browsers","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Implementation; Architecture; Reuse; Domain (mathematical analysis); Reference architecture; Web browser; World Wide Web; Software engineering; The Internet; Software architecture; Programming language; Engineering","score_opus":0.024059158381732226,"score_gpt":0.2762850183173075,"score_spread":0.2522258599355753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122010595","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017112924,0.0003709265,0.9659629,0.0006051367,0.00007799014,0.00017773247,0.0000973654,0.0049997726,0.01059524],"genre_scores_gemma":[0.1555669,0.0006352376,0.8267047,0.0003317428,0.0000639258,0.0003467271,0.00058917864,0.0017340655,0.014027564],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99657863,0.0010821156,0.00040172762,0.00049831584,0.0011316793,0.00030748727],"domain_scores_gemma":[0.9915667,0.0011203195,0.00034361263,0.003374419,0.003158872,0.00043599337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004736832,0.0005900014,0.0006932757,0.0020528643,0.0019754365,0.0053392258,0.0032271652,0.0038662436,0.004257798],"category_scores_gemma":[0.012311445,0.0011269056,0.0012721991,0.0017442892,0.0028358272,0.007003838,0.0035732244,0.0031959857,0.0025427588],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013055952,0.00012889234,0.0026383447,0.00024551406,0.00005017338,0.0007533091,0.0033076669,0.02471747,0.015107991,0.82444084,0.0087646,0.119714655],"study_design_scores_gemma":[0.00011398821,0.00033184234,0.0020621563,0.00062253233,0.00013376791,0.0019409428,0.0008519256,0.19703846,0.021358743,0.4010762,0.37425798,0.00021155614],"about_ca_topic_score_codex":0.0066588153,"about_ca_topic_score_gemma":0.0074009295,"teacher_disagreement_score":0.0066588153,"about_ca_system_score_codex":0.0017039592,"about_ca_system_score_gemma":0.0032014146,"threshold_uncertainty_score":0.025051057},"labels":[],"label_agreement":null},{"id":"W2122060876","doi":"10.5555/2337223.2337230","title":"Recovering traceability links between an API and its learning resources","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":125,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Traceability; Documentation; Ambiguity; Code (set theory); Source code; Context (archaeology); World Wide Web; Information retrieval; Software engineering; Programming language","score_opus":0.036843110419604065,"score_gpt":0.2864408961637425,"score_spread":0.24959778574413846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122060876","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22764398,0.0012243884,0.72069716,0.0009114249,0.00019689117,0.0004675168,0.0024593247,0.038781893,0.0076174885],"genre_scores_gemma":[0.55359197,0.0005341832,0.43083102,0.00019996722,0.0000804679,0.00022241475,0.007181194,0.0017713276,0.00558754],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9957231,0.0009924234,0.00052639295,0.00089948135,0.0016474137,0.00021104194],"domain_scores_gemma":[0.95587087,0.019381355,0.0056316345,0.009776589,0.008674102,0.0006654363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033957153,0.0011436086,0.00058548775,0.010702922,0.0014703514,0.0032255014,0.0018115607,0.0019855346,0.0023938252],"category_scores_gemma":[0.043250997,0.00058725773,0.0009096708,0.0046039224,0.00079640467,0.0053252424,0.0035515623,0.0020203786,0.0020414956],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003258266,0.00069710327,0.056361336,0.0007527024,0.00019402872,0.0006774426,0.0016042053,0.008342871,0.031068167,0.0060149385,0.013184664,0.88077676],"study_design_scores_gemma":[0.000098881166,0.00050560065,0.05587178,0.00058711704,0.00041089204,0.0023672113,0.0021771188,0.5364533,0.29501343,0.033215236,0.07306263,0.00023680442],"about_ca_topic_score_codex":0.010482926,"about_ca_topic_score_gemma":0.008976071,"teacher_disagreement_score":0.010702922,"about_ca_system_score_codex":0.0011680535,"about_ca_system_score_gemma":0.002495867,"threshold_uncertainty_score":0.020843863},"labels":[],"label_agreement":null},{"id":"W2122358298","doi":"10.1109/wpc.2003.1199195","title":"Identifying comprehension bottlenecks using program slicing and cognitive complexity metrics","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Program slicing; Computer science; Program comprehension; Software quality; Software engineering; Software maintenance; Software metric; Software construction; Software system; Slicing; Software; Static program analysis; Software development; Verification and validation; Quality (philosophy); Programming language; Engineering; World Wide Web","score_opus":0.12729209984737058,"score_gpt":0.3655375368444781,"score_spread":0.2382454369971075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122358298","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7859686,0.00051993626,0.20864645,0.00012659117,0.000015232568,0.00017365448,0.00073622173,0.0017084136,0.0021048957],"genre_scores_gemma":[0.9157189,0.00015529347,0.08220509,0.000014252996,0.000015148726,0.00015000487,0.0011942243,0.00017489745,0.00037222175],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99761033,0.0006245812,0.00031563462,0.00039218174,0.0008905193,0.00016674545],"domain_scores_gemma":[0.9522947,0.026993532,0.010132619,0.0026180262,0.006665676,0.0012953299],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039684833,0.001368845,0.0008824029,0.00662701,0.0004601028,0.0016908395,0.0007436455,0.0007586606,0.0012728834],"category_scores_gemma":[0.042477448,0.0003522075,0.0005218969,0.0032818487,0.0005462871,0.0046889205,0.0008947218,0.0006219483,0.00026423382],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009606618,0.0006749354,0.4328422,0.00066206633,0.0004197509,0.00026978538,0.0029464706,0.056067448,0.040257692,0.0050829183,0.0020437124,0.45777237],"study_design_scores_gemma":[0.00008748399,0.0015185311,0.39795783,0.00009352736,0.0002860499,0.00038853177,0.0010554888,0.54798275,0.03444141,0.013950227,0.002036322,0.00020172106],"about_ca_topic_score_codex":0.00740539,"about_ca_topic_score_gemma":0.00614841,"teacher_disagreement_score":0.00740539,"about_ca_system_score_codex":0.0007927083,"about_ca_system_score_gemma":0.0012872828,"threshold_uncertainty_score":0.02098757},"labels":[],"label_agreement":null},{"id":"W2122517552","doi":"10.5555/381473.381624","title":"2nd international workshop on living with inconsistency","year":2001,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Correctness; Computer science; Software engineering; Implementation; Focus (optics); Programming language; Software; Model checking","score_opus":0.032772682589832,"score_gpt":0.2788894805009732,"score_spread":0.24611679791114116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122517552","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009849258,0.0691444,0.46057034,0.07206611,0.09382128,0.0011640548,0.001600516,0.0051146443,0.2866694],"genre_scores_gemma":[0.05742067,0.046575807,0.2906639,0.013268382,0.023023121,0.0013538814,0.009072276,0.0039316383,0.55469036],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9954581,0.001498616,0.00038079242,0.00085802993,0.0013312774,0.000473054],"domain_scores_gemma":[0.99449724,0.0014634089,0.00016115654,0.0012492627,0.0017887584,0.0008402568],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006178676,0.001862588,0.0018377949,0.0020783946,0.0022972396,0.00752795,0.004258431,0.004173069,0.07168199],"category_scores_gemma":[0.011024597,0.00093911,0.0028688428,0.0021323804,0.001827561,0.01096513,0.007993461,0.005559908,0.02423796],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026696824,0.0002985936,0.0005280454,0.0008402152,0.00010840231,0.000605265,0.0012934741,0.0020640108,0.0027696986,0.06393534,0.5297355,0.3975543],"study_design_scores_gemma":[0.00004097846,0.00006603797,0.00034962012,0.0004766454,0.00005221707,0.00054809125,0.0004359593,0.0027339524,0.0011191301,0.04082681,0.9533105,0.000040062725],"about_ca_topic_score_codex":0.001955038,"about_ca_topic_score_gemma":0.0027751403,"teacher_disagreement_score":0.07168199,"about_ca_system_score_codex":0.0016299076,"about_ca_system_score_gemma":0.0024848743,"threshold_uncertainty_score":0.23980016},"labels":[],"label_agreement":null},{"id":"W2122525845","doi":"10.1109/icsm.2001.972708","title":"Supporting software maintenance by mining software update records","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software maintenance; Relevance (law); Context (archaeology); Software; Relation (database); Software engineering; Code (set theory); Software bug; Software system; Data mining; Programming language; Set (abstract data type)","score_opus":0.01790051301451302,"score_gpt":0.2614285986194306,"score_spread":0.24352808560491757,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122525845","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33775997,0.0017338162,0.6375974,0.0016100733,0.000087126275,0.0010729729,0.008257636,0.00587905,0.0060019144],"genre_scores_gemma":[0.40892786,0.0009840908,0.56780297,0.00012745288,0.00014453712,0.00046431704,0.020406082,0.00018075567,0.0009619409],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99439853,0.0016941383,0.0006207408,0.00084455527,0.0021682705,0.00027377205],"domain_scores_gemma":[0.95314205,0.032683488,0.004913778,0.0045471513,0.004306006,0.00040755345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00613989,0.00088664255,0.001146022,0.01037145,0.001330773,0.0028209798,0.0023296147,0.00117764,0.0012080732],"category_scores_gemma":[0.038364865,0.0006407211,0.0011397402,0.0074957027,0.0006538049,0.004326781,0.0017303836,0.0012533261,0.00113445],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036356493,0.0012083562,0.14784858,0.0019410767,0.0003122252,0.0007217936,0.0029431817,0.02894903,0.017017486,0.0067355544,0.0056745484,0.7862847],"study_design_scores_gemma":[0.00023181246,0.0008452389,0.10717444,0.0011325658,0.0009913773,0.0025364712,0.0046483264,0.66522014,0.09645066,0.05809173,0.06237054,0.0003067287],"about_ca_topic_score_codex":0.0034342029,"about_ca_topic_score_gemma":0.0066734785,"teacher_disagreement_score":0.01037145,"about_ca_system_score_codex":0.0006309385,"about_ca_system_score_gemma":0.0019703477,"threshold_uncertainty_score":0.03247124},"labels":[],"label_agreement":null},{"id":"W2122581326","doi":"10.1109/tse.2008.36","title":"Do Crosscutting Concerns Cause Defects?","year":2008,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":238,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Calgary","funders":"","keywords":"Computer science; Harm; Process (computing); Code (set theory); Measure (data warehouse); Source code; Degree (music); Software engineering; Risk analysis (engineering); Reliability engineering; Data mining; Programming language; Law; Engineering","score_opus":0.03323913553247372,"score_gpt":0.26905700605010585,"score_spread":0.23581787051763214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122581326","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9711368,0.0018717622,0.017798536,0.0022059586,0.00005769168,0.00007994906,0.0001806449,0.00018295707,0.006485745],"genre_scores_gemma":[0.99630094,0.00038871358,0.0025288034,0.00030754905,0.000024146399,0.00002635846,0.000084087646,0.000030138272,0.0003092915],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9887832,0.0032332009,0.00068864325,0.0019640469,0.0048028054,0.000528164],"domain_scores_gemma":[0.75856733,0.17880648,0.040913995,0.008663506,0.01160007,0.0014485464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008900899,0.00074954325,0.000618793,0.00284359,0.0006285419,0.0014615393,0.0010076741,0.002403595,0.0026320193],"category_scores_gemma":[0.10078901,0.0006390598,0.0007191138,0.002047856,0.0022937777,0.0041517345,0.0012760164,0.0013692583,0.00024735762],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004562408,0.00079616264,0.81277245,0.0013844214,0.0005794577,0.0016029058,0.005722561,0.002528745,0.025139961,0.012549148,0.0012016356,0.13526647],"study_design_scores_gemma":[0.000104294464,0.0014268291,0.88857555,0.000447707,0.000660256,0.0051280996,0.007437159,0.009449279,0.04469799,0.034781355,0.007200688,0.00009072414],"about_ca_topic_score_codex":0.0012998249,"about_ca_topic_score_gemma":0.0019906075,"teacher_disagreement_score":0.008900899,"about_ca_system_score_codex":0.00091396115,"about_ca_system_score_gemma":0.0008329307,"threshold_uncertainty_score":0.047073007},"labels":[],"label_agreement":null},{"id":"W2122649993","doi":"10.1109/issre.2001.989482","title":"Revisiting strategies for ordering class integration testing in the presence of dependency cycles","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Dependency (UML); Computer science; Dependency graph; Inheritance (genetic algorithm); Class (philosophy); Class diagram; Java; Theoretical computer science; Context (archaeology); Graph; Software; Programming language; Unified Modeling Language; Artificial intelligence","score_opus":0.04889506474204778,"score_gpt":0.31382897241151475,"score_spread":0.26493390766946695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122649993","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02504596,0.000488566,0.9674519,0.0007713113,0.00006141527,0.0005399044,0.000036239948,0.0010770386,0.004527746],"genre_scores_gemma":[0.14496286,0.00035674626,0.8521914,0.00024490114,0.000046221525,0.00024507404,0.000112193746,0.00022051268,0.0016201803],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99207616,0.0034012091,0.00086965103,0.0009203145,0.0022642608,0.00046841416],"domain_scores_gemma":[0.9659177,0.021810941,0.0022053644,0.0044038566,0.0048689223,0.00079316075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009942265,0.0019151561,0.0012187638,0.0040137,0.0010085012,0.0032539293,0.0030854999,0.0017902538,0.002438596],"category_scores_gemma":[0.026515406,0.00087224535,0.00088495156,0.0022041544,0.002336764,0.0045199813,0.0022150625,0.0021657764,0.0008344221],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030326511,0.00070848817,0.0055974,0.00062177854,0.000104119514,0.00089436845,0.002636599,0.055687424,0.027403947,0.11966536,0.0031868822,0.7831903],"study_design_scores_gemma":[0.0006369259,0.0013758525,0.0036598626,0.0007215059,0.00052606035,0.0022010847,0.0024668665,0.63618207,0.08064122,0.22460476,0.046673883,0.0003099736],"about_ca_topic_score_codex":0.0034788158,"about_ca_topic_score_gemma":0.0063939705,"teacher_disagreement_score":0.009942265,"about_ca_system_score_codex":0.0014158214,"about_ca_system_score_gemma":0.0034294948,"threshold_uncertainty_score":0.052580357},"labels":[],"label_agreement":null},{"id":"W2122718295","doi":"10.7287/peerj.preprints.826v1","title":"An empirical study of goto in C code","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Goto; Go/no go; Commit; Computer science; Statement (logic); Code (set theory); Empirical research; Dijkstra's algorithm; Programming language; Limit (mathematics); Statistics; Mathematics; Theoretical computer science; Machine learning; Law; Political science; Database; Set (abstract data type); Graph; Shortest path problem","score_opus":0.09214601773672157,"score_gpt":0.39471980448581767,"score_spread":0.3025737867490961,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122718295","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9956371,0.00018075533,0.0007133651,0.00035318814,0.00000916307,0.000059162212,0.00013235358,0.000026697162,0.0028883382],"genre_scores_gemma":[0.99768853,0.00014888938,0.00090737175,0.00014513786,0.000011612674,0.000050917395,0.00023066836,0.000045587196,0.0007711233],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98542494,0.005386125,0.0010256529,0.00169572,0.0055935206,0.0008740806],"domain_scores_gemma":[0.6235436,0.2614383,0.059109587,0.014926677,0.03541363,0.0055682887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011253402,0.00032390738,0.00031493598,0.003170518,0.0020316106,0.0024762505,0.0012658125,0.0011250505,0.0021706696],"category_scores_gemma":[0.17576379,0.0005431214,0.00026887655,0.003939009,0.003921851,0.0049894755,0.0019500674,0.0025119937,0.0005158586],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003929659,0.0005919943,0.90455854,0.00046189944,0.00006938513,0.00078958704,0.046479523,0.00044340696,0.0017100341,0.001977114,0.0020542115,0.040471457],"study_design_scores_gemma":[0.000040083716,0.00082339527,0.92378134,0.00046639398,0.000051848296,0.0011329413,0.05551569,0.00336977,0.002080125,0.0013379971,0.011305911,0.00009443397],"about_ca_topic_score_codex":0.00845245,"about_ca_topic_score_gemma":0.011551135,"teacher_disagreement_score":0.011253402,"about_ca_system_score_codex":0.0016390551,"about_ca_system_score_gemma":0.0017038301,"threshold_uncertainty_score":0.059514403},"labels":[],"label_agreement":null},{"id":"W2122723127","doi":"10.1016/j.scico.2010.10.006","title":"Metamodeling semantics of multiple inheritance","year":2010,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Inheritance (genetic algorithm); Metamodeling; Programming language; Programmer; Semantics (computer science); Context (archaeology); Construct (python library); Reusability; Multiple inheritance; Theoretical computer science; Object-oriented programming","score_opus":0.019897740592626993,"score_gpt":0.27734212334047276,"score_spread":0.25744438274784576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122723127","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02199855,0.00045601834,0.9654237,0.00095763523,0.00013931416,0.000053667027,0.0001405873,0.0012750077,0.009555483],"genre_scores_gemma":[0.49537605,0.000790171,0.4943274,0.0005289816,0.00021942829,0.0001748059,0.00045320063,0.00074376474,0.007386119],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99751246,0.00085877633,0.00034814043,0.00036766485,0.0006827975,0.00023007386],"domain_scores_gemma":[0.9966941,0.0011932908,0.00026816892,0.0009814068,0.00069017935,0.0001729031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004195493,0.00058549613,0.00074881467,0.0012801704,0.0014598179,0.004237384,0.002279621,0.0016304874,0.0025511405],"category_scores_gemma":[0.005476007,0.00094098557,0.0015377598,0.0011857423,0.002633816,0.0091592,0.0025952084,0.0030169094,0.0005279492],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026787508,0.000025771664,0.0002294967,0.000054269673,0.000014345925,0.000081883896,0.0006802114,0.0025240595,0.0015795153,0.981427,0.0006464408,0.012710276],"study_design_scores_gemma":[0.000032971606,0.000018980207,0.00013137987,0.00007329357,0.00005258075,0.0001523465,0.00015752405,0.019959744,0.003495107,0.9554601,0.02043901,0.000026813852],"about_ca_topic_score_codex":0.002537595,"about_ca_topic_score_gemma":0.0032169123,"teacher_disagreement_score":0.004237384,"about_ca_system_score_codex":0.0018707331,"about_ca_system_score_gemma":0.0017429049,"threshold_uncertainty_score":0.022188187},"labels":[],"label_agreement":null},{"id":"W2122744555","doi":"10.1109/ase.2008.81","title":"Semi-Automating Pragmatic Reuse Tasks","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Computer science; Task (project management); Scope (computer science); Software engineering; Process (computing); Code reuse; Plan (archaeology); Human–computer interaction; Source code; Code (set theory); Systems engineering; Programming language; Software; Engineering; Set (abstract data type)","score_opus":0.019888137130382424,"score_gpt":0.26197083371904767,"score_spread":0.24208269658866524,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122744555","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016267411,0.00010327995,0.96743983,0.00025256688,0.000060097715,0.00048017292,0.00014824959,0.012939534,0.002308898],"genre_scores_gemma":[0.08940712,0.00010818165,0.9045109,0.00015079959,0.000030835974,0.00045461708,0.000741243,0.0021780152,0.00241838],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96774065,0.015676828,0.0028231721,0.004131054,0.008236592,0.0013916597],"domain_scores_gemma":[0.8759355,0.07329516,0.0057619037,0.032558303,0.010931299,0.0015178791],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017169496,0.002985783,0.0018901583,0.002174571,0.002018912,0.004645714,0.0039635766,0.0025955008,0.005150785],"category_scores_gemma":[0.09669792,0.0025567366,0.0024729301,0.0011766349,0.0026648003,0.0057305237,0.0111209955,0.0039660884,0.0038433375],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007170152,0.00080888957,0.005786538,0.0023077535,0.00027185344,0.0011203111,0.015662951,0.030568147,0.090458795,0.041480314,0.018988328,0.7918291],"study_design_scores_gemma":[0.000685849,0.00071641215,0.00693211,0.00097213476,0.00040265764,0.002394006,0.0061603473,0.46523926,0.16211137,0.17669708,0.17685378,0.0008350914],"about_ca_topic_score_codex":0.0042405464,"about_ca_topic_score_gemma":0.006194801,"teacher_disagreement_score":0.017169496,"about_ca_system_score_codex":0.0011800623,"about_ca_system_score_gemma":0.0072127683,"threshold_uncertainty_score":0.090802014},"labels":[],"label_agreement":null},{"id":"W2123187470","doi":"10.1109/ase.2003.1240310","title":"Automatically inferring concern code from program investigation activities","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; University of British Columbia","keywords":"Computer science; Source code; Session (web analytics); Task (project management); Program comprehension; Code (set theory); Context (archaeology); Set (abstract data type); Programming language; Program analysis; Program code; Object (grammar); Software engineering; Software; World Wide Web; Artificial intelligence; Software system; Engineering","score_opus":0.037618740272692575,"score_gpt":0.29942745658008957,"score_spread":0.261808716307397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2123187470","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6297134,0.00027309512,0.36373276,0.00030306843,0.000025982215,0.00038564095,0.0013869901,0.0026684385,0.0015106136],"genre_scores_gemma":[0.69347674,0.00021533476,0.29845843,0.00008000248,0.000032428317,0.00042468117,0.0055729034,0.00024870163,0.0014907676],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9955954,0.0020015466,0.00029702947,0.00082330627,0.0010698187,0.00021284418],"domain_scores_gemma":[0.9385136,0.044462666,0.006143129,0.0038242077,0.006499647,0.00055685634],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033091907,0.00086538336,0.00058723957,0.0040084994,0.0005873879,0.0010663647,0.00091612514,0.0012896832,0.0007402061],"category_scores_gemma":[0.041892834,0.00060999556,0.00059247896,0.001629633,0.00052529114,0.0015203361,0.0010674079,0.0011018866,0.00058302836],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089826365,0.0006995461,0.24096285,0.0010592155,0.0001470535,0.0011457963,0.01075326,0.011025965,0.106983006,0.0027935025,0.0044346116,0.619097],"study_design_scores_gemma":[0.00013865206,0.0008075072,0.33215693,0.00026810027,0.00028570235,0.0025823785,0.006186023,0.51969945,0.10762587,0.010185676,0.019858865,0.00020479342],"about_ca_topic_score_codex":0.005528782,"about_ca_topic_score_gemma":0.00962514,"teacher_disagreement_score":0.005528782,"about_ca_system_score_codex":0.0006144913,"about_ca_system_score_gemma":0.0014034042,"threshold_uncertainty_score":0.017500877},"labels":[],"label_agreement":null},{"id":"W2123252005","doi":"10.1109/icdar.2001.953813","title":"On-line recognition of UML diagrams","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; UML tool; Applications of UML; Unified Modeling Language; Class diagram; Systems Modeling Language; Programming language; Software; Software engineering","score_opus":0.0674277720505536,"score_gpt":0.2760105024533249,"score_spread":0.2085827304027713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2123252005","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012280878,0.00012292588,0.92162126,0.0001253437,0.00011367528,0.0002669104,0.0009069906,0.053561144,0.0110009145],"genre_scores_gemma":[0.14573061,0.00032420942,0.8139382,0.00023806942,0.00009785988,0.00033266327,0.005624443,0.0104384115,0.023275457],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99708635,0.00068874273,0.00024012149,0.00066554744,0.0011193929,0.000199868],"domain_scores_gemma":[0.9895216,0.00483752,0.00063098955,0.0027561937,0.0019353492,0.0003183121],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015976485,0.0011980549,0.0009326587,0.0023525693,0.0005685242,0.0034984779,0.0018284633,0.0011625604,0.023588188],"category_scores_gemma":[0.013510791,0.00054387696,0.0008825357,0.00086630415,0.00045472625,0.0026786905,0.0017960619,0.0012586687,0.013642801],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006112371,0.00023481925,0.0017844925,0.00036331918,0.000039948,0.0006906844,0.0018474183,0.0036564865,0.07352122,0.009709282,0.030533174,0.877008],"study_design_scores_gemma":[0.0001551707,0.00027956764,0.0059956633,0.00039366898,0.00009036995,0.0027417324,0.0010403376,0.29420334,0.39228705,0.021450572,0.28117973,0.00018287222],"about_ca_topic_score_codex":0.0015219796,"about_ca_topic_score_gemma":0.001640105,"teacher_disagreement_score":0.023588188,"about_ca_system_score_codex":0.0007550422,"about_ca_system_score_gemma":0.00075180165,"threshold_uncertainty_score":0.07891035},"labels":[],"label_agreement":null},{"id":"W2123316933","doi":"10.1007/s00766-009-0088-6","title":"Developing comprehensive acceptance tests from use cases and robustness diagrams","year":2009,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"King Fahd University of Petroleum and Minerals","keywords":"Robustness (evolution); Robustness testing; Computer science; Acceptance testing; Reliability engineering; Software engineering; Engineering; Programming language; Software","score_opus":0.07690407596900042,"score_gpt":0.30633885552576606,"score_spread":0.22943477955676564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2123316933","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.120081864,0.00020601237,0.8584043,0.00024811496,0.00003609199,0.001018033,0.00071409385,0.015374994,0.003916535],"genre_scores_gemma":[0.42540538,0.00014384494,0.567302,0.00014034497,0.000023459748,0.0008091228,0.002673923,0.0018893578,0.0016125954],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98355955,0.0065574097,0.0013329741,0.0012937988,0.006340775,0.0009155409],"domain_scores_gemma":[0.8460724,0.12118455,0.006483313,0.010250767,0.015048644,0.0009603098],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010859839,0.002418653,0.001355634,0.004880234,0.00052033446,0.0022902645,0.0028513612,0.0024208562,0.004429192],"category_scores_gemma":[0.1079705,0.0016962278,0.002448216,0.0013949549,0.0011279285,0.0051604263,0.00292474,0.002258201,0.0017889813],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008758894,0.0016456876,0.02350394,0.0021079131,0.0007668271,0.0026018296,0.0027116623,0.26629046,0.08498642,0.02368848,0.00750201,0.5833188],"study_design_scores_gemma":[0.00021384825,0.00075314794,0.0054954384,0.00036133575,0.00031973183,0.00070682203,0.00063851546,0.9182461,0.048017718,0.018404286,0.0067262687,0.00011684216],"about_ca_topic_score_codex":0.0034285984,"about_ca_topic_score_gemma":0.005354403,"teacher_disagreement_score":0.010859839,"about_ca_system_score_codex":0.0010450259,"about_ca_system_score_gemma":0.0022235487,"threshold_uncertainty_score":0.05743295},"labels":[],"label_agreement":null},{"id":"W2123980061","doi":"10.1109/icsm.2001.972738","title":"Integrating information sources for visualizing Java programs","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"University of Victoria","keywords":"Computer science; Java; Documentation; Visualization; Source code; Presentation (obstetrics); Software engineering; Internal documentation; Information integration; Domain (mathematical analysis); Software visualization; Software; World Wide Web; Programming language; Database; Software development; Data mining; Component-based software engineering; Software construction","score_opus":0.03635503188684518,"score_gpt":0.2832856276809855,"score_spread":0.24693059579414028,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2123980061","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010162794,0.0013246966,0.9479798,0.0007597399,0.00012381283,0.0002754016,0.0016382016,0.029971685,0.0077638975],"genre_scores_gemma":[0.061923575,0.0016880856,0.9248073,0.00016667592,0.000102816775,0.0004346921,0.0033943632,0.0048474134,0.0026351348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99715143,0.00085564196,0.00033979223,0.00028966155,0.0012202628,0.00014314937],"domain_scores_gemma":[0.99316484,0.00358733,0.0003934326,0.0012211008,0.0013459545,0.00028741185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003276173,0.001820352,0.0011678332,0.0110023115,0.0013332207,0.0067838477,0.0015027269,0.0021227824,0.008482467],"category_scores_gemma":[0.014275263,0.0012115482,0.0016923099,0.008445308,0.0006679721,0.0076289927,0.0053276145,0.002138721,0.0023406034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082759844,0.00041216688,0.0027563483,0.0021653052,0.00030739672,0.0011214216,0.008300768,0.008776493,0.03230979,0.048799783,0.028000271,0.8662226],"study_design_scores_gemma":[0.00047077987,0.00034937527,0.0059079803,0.0021257708,0.0009620854,0.0021827398,0.0025701746,0.17021362,0.13073282,0.13482389,0.5487167,0.0009440512],"about_ca_topic_score_codex":0.0025368023,"about_ca_topic_score_gemma":0.0034157222,"teacher_disagreement_score":0.0110023115,"about_ca_system_score_codex":0.000717097,"about_ca_system_score_gemma":0.0011456225,"threshold_uncertainty_score":0.028376639},"labels":[],"label_agreement":null},{"id":"W2124297776","doi":"10.1109/seaa.2008.75","title":"IFPUG-COSMIC Statistical Conversion","year":2008,"lang":"en","type":"article","venue":"Proceedings of the ... EUROMICRO Conference/EUROMICRO","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Computer science; Verifiable secret sharing; Process (computing); COSMIC cancer database; Function (biology); Data collection; Quality (philosophy); Interval (graph theory); Data mining; Data science; Set (abstract data type); Mathematics; Statistics; Programming language; Astronomy","score_opus":0.02702442971313656,"score_gpt":0.23003430389491566,"score_spread":0.2030098741817791,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124297776","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1074245,0.00088210707,0.7859171,0.00070860947,0.00070537475,0.0008012936,0.030683061,0.02975217,0.0431258],"genre_scores_gemma":[0.4187405,0.00048087805,0.53020006,0.0004085271,0.00023267385,0.0027436318,0.0331903,0.0060894806,0.007913922],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99109215,0.0020564154,0.0007676356,0.0017267264,0.0038509318,0.0005061744],"domain_scores_gemma":[0.9723479,0.008254258,0.0013651146,0.012389248,0.0053725136,0.00027088725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006218089,0.0011409905,0.001147396,0.008294187,0.0011571187,0.003926162,0.0019133723,0.0010271339,0.011733691],"category_scores_gemma":[0.051636267,0.0005231433,0.0012847676,0.009459124,0.0015471221,0.0025220767,0.0034792554,0.0019096797,0.004159146],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007426031,0.00020575583,0.047822446,0.00089318654,0.00024166411,0.0006577512,0.0017116135,0.021114102,0.011608498,0.16253668,0.058597326,0.69386834],"study_design_scores_gemma":[0.00011173423,0.00057115115,0.16198502,0.0006089691,0.00018974088,0.002926156,0.0016191824,0.15139341,0.068946384,0.19564223,0.4156377,0.00036827652],"about_ca_topic_score_codex":0.004319365,"about_ca_topic_score_gemma":0.0023892915,"teacher_disagreement_score":0.011733691,"about_ca_system_score_codex":0.001620097,"about_ca_system_score_gemma":0.0015061923,"threshold_uncertainty_score":0.039253056},"labels":[],"label_agreement":null},{"id":"W2124312992","doi":"10.1109/csse.2008.1015","title":"An Algorithm of System Decomposition Based on Laplace Spectral Graph Partitioning Technology","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Graph partition; Graph; Computer science; Strength of a graph; Spectral graph theory; Voltage graph; Null graph; Algebraic graph theory; Laplace transform; Algorithm; Graph theory; Graph bandwidth; Decomposition; Mathematics; Theoretical computer science; Line graph; Combinatorics","score_opus":0.009852692859426167,"score_gpt":0.2567739556670001,"score_spread":0.24692126280757395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124312992","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015267784,0.00002326937,0.9975744,0.000020745609,0.000010053344,0.000022772687,0.000009766996,0.00033836023,0.0004738771],"genre_scores_gemma":[0.06749139,0.000075898504,0.93075156,0.000051361087,0.00002240917,0.00016047241,0.00016309293,0.0001817392,0.0011020107],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995018,0.00010806653,0.000032525964,0.00011634529,0.00018360502,0.00005766715],"domain_scores_gemma":[0.99959165,0.00015187333,0.000027296002,0.000049187798,0.00016024936,0.000019771527],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00052374374,0.0010063152,0.0006726528,0.0016947759,0.0007328783,0.0007099027,0.00082420366,0.0006490447,0.003831797],"category_scores_gemma":[0.0013916634,0.0003749057,0.0007717027,0.00097154366,0.00062019704,0.0014187531,0.0010325168,0.0008860134,0.0013725576],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014684361,0.000104175684,0.00089618965,0.00024186693,0.00007108772,0.00016762054,0.00035044173,0.17148809,0.034444522,0.07488692,0.006290212,0.710912],"study_design_scores_gemma":[0.000038442093,0.00004540692,0.0002131591,0.000013187151,0.00002437679,0.00012837675,0.000050485898,0.95474696,0.010880844,0.029047014,0.0047932616,0.000018611874],"about_ca_topic_score_codex":0.0021213,"about_ca_topic_score_gemma":0.0016546431,"teacher_disagreement_score":0.003831797,"about_ca_system_score_codex":0.0006549272,"about_ca_system_score_gemma":0.00088710774,"threshold_uncertainty_score":0.0128186345},"labels":[],"label_agreement":null},{"id":"W2124518380","doi":"10.1109/apsec.2000.896710","title":"A comparative evaluation of techniques for syntactic level source code analysis","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Parsing; Computer science; Programming language; Syntax; Source code; Abstract syntax tree; Code (set theory); Software; Artificial intelligence; Simple (philosophy); Natural language processing; Abstract syntax; Set (abstract data type)","score_opus":0.21296336926181794,"score_gpt":0.39342455532694104,"score_spread":0.1804611860651231,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124518380","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19841217,0.05164195,0.6766661,0.002108339,0.0008782764,0.0016127839,0.0032430985,0.03847303,0.026964258],"genre_scores_gemma":[0.28527048,0.022127036,0.6717298,0.0005015834,0.0003849149,0.0007252343,0.008615077,0.0051617827,0.0054840483],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9680862,0.010558807,0.0024056947,0.003200852,0.014650418,0.0010980343],"domain_scores_gemma":[0.8958637,0.075282566,0.002862126,0.009604579,0.015409159,0.0009778589],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015559716,0.002233337,0.0016421737,0.011042293,0.0010425906,0.0030832218,0.004235942,0.0031650686,0.0032297217],"category_scores_gemma":[0.052630607,0.0010821775,0.0025667408,0.010379901,0.0012200988,0.0070095886,0.003172214,0.0019682008,0.0024948556],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004061442,0.00085564447,0.006725061,0.005169124,0.0012809978,0.00030126295,0.0016662298,0.007985872,0.026626965,0.006960003,0.009240991,0.92912656],"study_design_scores_gemma":[0.0026414872,0.010473291,0.09504941,0.004043081,0.005374645,0.0066899867,0.0072426083,0.4988413,0.2073684,0.035153095,0.12582111,0.0013015694],"about_ca_topic_score_codex":0.002894575,"about_ca_topic_score_gemma":0.0039001114,"teacher_disagreement_score":0.015559716,"about_ca_system_score_codex":0.0011299617,"about_ca_system_score_gemma":0.0017990952,"threshold_uncertainty_score":0.08228862},"labels":[],"label_agreement":null},{"id":"W2124740781","doi":"10.1145/986710.986720","title":"The programmer life-cycle","year":2004,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Programmer; Productivity; Computer science; Construct (python library); Software; Programming language; Economics; Macroeconomics","score_opus":0.0147598033736386,"score_gpt":0.24609978333871485,"score_spread":0.23133997996507624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124740781","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30248126,0.0152841415,0.22295517,0.016303679,0.0005849268,0.00093312,0.0068074972,0.0029456867,0.4317044],"genre_scores_gemma":[0.8642478,0.0066246823,0.039035153,0.0022255722,0.00028469588,0.0007060672,0.0027630893,0.0009569976,0.083155945],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9967277,0.00083800306,0.00013879374,0.00047894643,0.0014218404,0.00039475397],"domain_scores_gemma":[0.9877945,0.0037614282,0.0013017819,0.0015325756,0.0038672034,0.0017424686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028055736,0.0005966071,0.000372874,0.0024085692,0.0012355708,0.0048276503,0.0012818597,0.0012364762,0.009975787],"category_scores_gemma":[0.02116797,0.0005301688,0.0004144281,0.0022208062,0.0012021787,0.0041327905,0.0023375745,0.0014566086,0.004672078],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004350451,0.00031274193,0.06769086,0.0006434175,0.00011785621,0.00044749858,0.006213661,0.013518359,0.0045519285,0.35700053,0.052596774,0.49647135],"study_design_scores_gemma":[0.00006592249,0.000502951,0.06591829,0.0005791332,0.00008658781,0.0014363584,0.0033610181,0.023174565,0.0032239845,0.251861,0.64962214,0.00016795429],"about_ca_topic_score_codex":0.0049615246,"about_ca_topic_score_gemma":0.0027462817,"teacher_disagreement_score":0.009975787,"about_ca_system_score_codex":0.0029207366,"about_ca_system_score_gemma":0.004625587,"threshold_uncertainty_score":0.033372343},"labels":[],"label_agreement":null},{"id":"W2124777132","doi":"10.1109/snpd.2015.7176280","title":"Systematic mapping study of missing values techniques in software engineering data","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Data mining; Missing data; Imputation (statistics); Software; Data science; Machine learning","score_opus":0.08569349859442171,"score_gpt":0.3158034108359187,"score_spread":0.23010991224149696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124777132","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3595761,0.225378,0.39314502,0.0026684378,0.0007913497,0.0068900916,0.006061318,0.00044608617,0.0050435485],"genre_scores_gemma":[0.69085574,0.047974247,0.25024801,0.00069038174,0.00017201537,0.005503014,0.0038378378,0.00010575015,0.000613002],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.8829787,0.067859344,0.024181042,0.00679039,0.017420331,0.0007702849],"domain_scores_gemma":[0.4372273,0.46121734,0.047351655,0.021518903,0.031829163,0.0008555996],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09158626,0.0010721596,0.0027664842,0.033427723,0.0014363554,0.0033958184,0.001867934,0.0013125851,0.001314821],"category_scores_gemma":[0.2946755,0.0008627837,0.0048505836,0.0345844,0.0015900854,0.0049122046,0.0031649605,0.0013122339,0.0002205069],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005867824,0.00035761323,0.19037434,0.11528321,0.013875633,0.0011365836,0.013509936,0.0042961477,0.004132227,0.011501003,0.0022757929,0.6426707],"study_design_scores_gemma":[0.00091945514,0.006670764,0.38635755,0.19166146,0.06694573,0.0071002506,0.06703015,0.04148486,0.04876902,0.06597788,0.11616655,0.0009163869],"about_ca_topic_score_codex":0.0013543551,"about_ca_topic_score_gemma":0.0023592343,"teacher_disagreement_score":0.90841377,"about_ca_system_score_codex":0.0017049602,"about_ca_system_score_gemma":0.010048436,"threshold_uncertainty_score":0.48436022},"labels":[],"label_agreement":null},{"id":"W2124961859","doi":"10.1109/icpc.2007.6","title":"A Comparative Study of Three Program Exploration Tools","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; University of British Columbia","funders":"","keywords":"Computer science; Task (project management); Variety (cybernetics); Software engineering; Software; Plan (archaeology); Human–computer interaction; Code (set theory); Programming language; Artificial intelligence; Systems engineering; Engineering","score_opus":0.15362414765086374,"score_gpt":0.38234565943176013,"score_spread":0.2287215117808964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124961859","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99756336,0.00011934099,0.0011937607,0.000022311478,0.0000038146056,0.0001815422,0.000035693785,0.00005645022,0.0008237623],"genre_scores_gemma":[0.99069864,0.00018158725,0.0077776466,0.000048497837,0.000007931873,0.00046139405,0.0002150301,0.000034092434,0.0005751],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9943461,0.0027964048,0.000733858,0.0005566082,0.0010079507,0.0005591049],"domain_scores_gemma":[0.9375277,0.04755719,0.0038917533,0.0033163435,0.004347104,0.0033598703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006090919,0.000715998,0.00074713497,0.0027240664,0.00093072705,0.0016593096,0.0012094756,0.0010415496,0.0015709125],"category_scores_gemma":[0.04875923,0.00043656115,0.00043403573,0.0017974337,0.0012736947,0.002738542,0.0018161201,0.0007610745,0.00025639532],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.03049041,0.052337922,0.15080701,0.005912224,0.0007987366,0.0025689774,0.10001727,0.005623526,0.12348111,0.0035520655,0.0021559792,0.5222548],"study_design_scores_gemma":[0.0054369774,0.18148088,0.65053,0.00058087136,0.0009804963,0.0034673724,0.04157172,0.01689714,0.070657365,0.0043308525,0.023529252,0.00053702097],"about_ca_topic_score_codex":0.00092798995,"about_ca_topic_score_gemma":0.0018663885,"teacher_disagreement_score":0.006090919,"about_ca_system_score_codex":0.0008569362,"about_ca_system_score_gemma":0.0010610006,"threshold_uncertainty_score":0.032212257},"labels":[],"label_agreement":null},{"id":"W2125031115","doi":"10.1145/1062455.1062617","title":"Predictor models in software engineering (PROMISE)","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Software engineering; Software; Data science; Systems engineering; Engineering; Programming language","score_opus":0.015097902808183327,"score_gpt":0.23018033647733424,"score_spread":0.21508243366915092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125031115","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014971172,0.027478816,0.8871693,0.030407578,0.0035500138,0.00011335828,0.004726858,0.00786424,0.02371867],"genre_scores_gemma":[0.34780025,0.050278515,0.46126035,0.003807276,0.010525277,0.000977932,0.0124432305,0.00339839,0.10950882],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973412,0.0017395902,0.00008339838,0.00030656054,0.00041752157,0.000111772446],"domain_scores_gemma":[0.98250455,0.012116775,0.0006950163,0.002174187,0.0019763133,0.00053312746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076811137,0.0011112515,0.0022501007,0.0014670898,0.0011885631,0.003141166,0.0020153143,0.0016106042,0.028311448],"category_scores_gemma":[0.023778362,0.0011284675,0.0013762382,0.004656588,0.000957444,0.0047205733,0.0027451995,0.0058773872,0.013996986],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005198081,0.00028157455,0.005363786,0.0013639138,0.00041041142,0.00014631869,0.00023614179,0.080813,0.0003770163,0.18502624,0.26238486,0.4630769],"study_design_scores_gemma":[0.00022159914,0.00034893904,0.003949202,0.00068025023,0.00044639004,0.00016183637,0.00013880021,0.38326743,0.0011576208,0.49348748,0.1160053,0.00013516894],"about_ca_topic_score_codex":0.00643552,"about_ca_topic_score_gemma":0.008514277,"teacher_disagreement_score":0.028311448,"about_ca_system_score_codex":0.00075152767,"about_ca_system_score_gemma":0.0032900197,"threshold_uncertainty_score":0.094711244},"labels":[],"label_agreement":null},{"id":"W2125181275","doi":"10.1109/hpcc.2011.60","title":"Social Network Analysis in Software Testing to Categorize Unit Test Cases Based on Coverage Information","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Test suite; Computer science; Test Management Approach; Regression testing; Test script; Test harness; Reliability engineering; Software quality; Unit testing; Test case; System under test; Software engineering; Software maintenance; Software; Software system; Software construction; Software development; Engineering; Operating system; Machine learning","score_opus":0.04989803146578312,"score_gpt":0.27172205451214154,"score_spread":0.22182402304635843,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125181275","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45199493,0.0012809978,0.53777766,0.00094211183,0.00004758619,0.00052400876,0.0014656187,0.00072069286,0.005246386],"genre_scores_gemma":[0.8090305,0.0004243813,0.18713804,0.00007826844,0.00007489182,0.00042440437,0.001891994,0.000066975794,0.00087047246],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960756,0.0020998204,0.00028677654,0.00060412893,0.0007806711,0.00015298644],"domain_scores_gemma":[0.97268933,0.021147527,0.0027039624,0.0013942866,0.0016391963,0.0004257075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029240309,0.00077940023,0.0008698094,0.015899705,0.0011156315,0.0014222222,0.0007571035,0.0011638253,0.001316114],"category_scores_gemma":[0.020265447,0.00036555846,0.00084106723,0.0066580866,0.0010343746,0.0026917295,0.001045989,0.00062383525,0.00031920086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008226163,0.0007827049,0.2554253,0.0011686519,0.0011130135,0.0009829367,0.0054959077,0.1171887,0.014693865,0.04600797,0.005452982,0.5508654],"study_design_scores_gemma":[0.000040256647,0.00016638055,0.07057127,0.00014580636,0.00025441914,0.0006670577,0.0012467753,0.8756443,0.0048117577,0.04129814,0.005062116,0.000091691305],"about_ca_topic_score_codex":0.0050078486,"about_ca_topic_score_gemma":0.0052745594,"teacher_disagreement_score":0.015899705,"about_ca_system_score_codex":0.0012722779,"about_ca_system_score_gemma":0.0005241542,"threshold_uncertainty_score":0.015463889},"labels":[],"label_agreement":null},{"id":"W2125220238","doi":"10.1109/services-i.2009.106","title":"Web-FIM: Automated Framework for the Inference of Business Software Models","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Computer Research Institute of Montréal","funders":"","keywords":"Computer science; Inference; Visualization; The Internet; Web service; Model checking; Software deployment; Software; Web application; Software engineering; Programming language; Data mining; World Wide Web; Artificial intelligence","score_opus":0.03666594797359515,"score_gpt":0.30921899417407844,"score_spread":0.27255304620048326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125220238","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008656772,0.000054037373,0.9774573,0.00011002566,0.000020534935,0.00008933454,0.00045992344,0.020181935,0.00076124474],"genre_scores_gemma":[0.048345905,0.00017246313,0.94488674,0.00016332748,0.000045551285,0.00043036763,0.0026365465,0.0019180478,0.0014011255],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99575484,0.001140244,0.0003904336,0.00074668246,0.0016499929,0.00031787177],"domain_scores_gemma":[0.9939568,0.0034962874,0.00038008415,0.0015014017,0.0005478223,0.00011761838],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051180054,0.0019059542,0.0014620472,0.0050076083,0.0013556703,0.0036531668,0.0047350903,0.001983201,0.010948914],"category_scores_gemma":[0.019374136,0.00174622,0.0044462746,0.0016151144,0.0019151033,0.0044821547,0.0039399406,0.003298891,0.0028537633],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003634457,0.00051335455,0.00430116,0.0013547268,0.00053952436,0.0020998055,0.00090921234,0.26870856,0.010621728,0.31655303,0.030304449,0.36373094],"study_design_scores_gemma":[0.000071881346,0.000054813096,0.00042464412,0.00019589529,0.00006901465,0.0004087807,0.0000679555,0.8162364,0.009082444,0.14384012,0.0294714,0.00007672203],"about_ca_topic_score_codex":0.014578994,"about_ca_topic_score_gemma":0.015492305,"teacher_disagreement_score":0.014578994,"about_ca_system_score_codex":0.0023935542,"about_ca_system_score_gemma":0.0035527637,"threshold_uncertainty_score":0.03662783},"labels":[],"label_agreement":null},{"id":"W2125279207","doi":"10.1109/icsm.2015.7332457","title":"Investigating code review quality: Do people and participation matter?","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":149,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université de Montréal; University of Waterloo","funders":"","keywords":"Code review; Computer science; Software quality; Quality (philosophy); Code (set theory); Software bug; Software engineering; Process (computing); Source code; KPI-driven code analysis; Empirical research; Static program analysis; Software inspection; Set (abstract data type); Software development; Software; Programming language","score_opus":0.11815631206648786,"score_gpt":0.38325071342079636,"score_spread":0.2650944013543085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125279207","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96994597,0.0021912288,0.015020751,0.003365429,0.00013801573,0.00039762747,0.00022170288,0.00017245831,0.008546804],"genre_scores_gemma":[0.99674726,0.00021745433,0.0016890072,0.00035715822,0.000077683835,0.00018134023,0.00010346693,0.00004667874,0.00057992106],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.87200934,0.07499455,0.009861512,0.009546622,0.029774303,0.0038136842],"domain_scores_gemma":[0.32596385,0.49200076,0.11411091,0.022587579,0.03530079,0.0100361025],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08487585,0.00054802303,0.0013317309,0.0048924903,0.0017353328,0.004750991,0.0014867783,0.0018916442,0.0037655511],"category_scores_gemma":[0.38878965,0.000638459,0.000926068,0.0038551714,0.0031781949,0.00672141,0.0033640247,0.0013646043,0.00088667],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042645333,0.0003232576,0.8890349,0.00054658944,0.00046380132,0.00020485811,0.026975231,0.0002792622,0.0007830325,0.0008402374,0.0012754684,0.078846864],"study_design_scores_gemma":[0.00010529494,0.0006545314,0.9637469,0.0004650004,0.00023622868,0.00056252204,0.01868034,0.0032719267,0.0011371821,0.003869693,0.007161692,0.000108575834],"about_ca_topic_score_codex":0.0021177772,"about_ca_topic_score_gemma":0.0027089925,"teacher_disagreement_score":0.9151242,"about_ca_system_score_codex":0.0015632249,"about_ca_system_score_gemma":0.0025231314,"threshold_uncertainty_score":0.4488718},"labels":[],"label_agreement":null},{"id":"W2125299728","doi":"10.1109/wi-iat.2011.196","title":"Dependency and Entropy Based Impact Analysis for Service-Oriented System Evolution","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Dependency (UML); Service-oriented architecture; Entropy (arrow of time); Information system; Pace; Data mining; Data science; Software engineering; Web service; Engineering; World Wide Web","score_opus":0.021772411460436115,"score_gpt":0.25763166230726,"score_spread":0.23585925084682385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125299728","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08041839,0.00030324658,0.911695,0.00021988059,0.00003151466,0.00009447472,0.00022365546,0.0002807942,0.0067329803],"genre_scores_gemma":[0.9194033,0.0003366419,0.07841362,0.00004698719,0.00006298515,0.00021320082,0.0002917313,0.00009462191,0.001136908],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99813724,0.0007109739,0.00009165799,0.00012558594,0.00082540984,0.00010918007],"domain_scores_gemma":[0.9915911,0.0065505574,0.00064430246,0.0004493103,0.00063955924,0.00012499512],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002290038,0.00067860686,0.0005811894,0.004422325,0.00054452533,0.00081767177,0.00050331827,0.00044655858,0.0014801268],"category_scores_gemma":[0.010245826,0.0003011789,0.0011232493,0.0021233552,0.001124999,0.0017996836,0.0010282506,0.00070688175,0.00013778333],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010822732,0.00009518446,0.00895157,0.000119202625,0.00015250291,0.0001834069,0.00017248851,0.78346217,0.0058665923,0.14771025,0.00060109945,0.0525773],"study_design_scores_gemma":[0.0000032153407,0.000021775564,0.0030278948,0.000008580475,0.000019991225,0.00003491186,0.000014112894,0.9544616,0.0013869094,0.04057129,0.00043142127,0.000018234638],"about_ca_topic_score_codex":0.0023848733,"about_ca_topic_score_gemma":0.0013678722,"teacher_disagreement_score":0.004422325,"about_ca_system_score_codex":0.0016611068,"about_ca_system_score_gemma":0.00076061627,"threshold_uncertainty_score":0.012111008},"labels":[],"label_agreement":null},{"id":"W2125343508","doi":"10.1007/978-3-642-23716-4_22","title":"A Fast Algorithm to Locate Concepts in Execution Traces","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Dynamic programming; TRACE (psycholinguistics); Genetic programming; Algorithm; Identification (biology); Genetic algorithm; Segmentation; Theoretical computer science; Artificial intelligence; Machine learning","score_opus":0.02557422175582425,"score_gpt":0.2716429627880335,"score_spread":0.24606874103220927,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125343508","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041139806,0.0002062934,0.9849508,0.000074553165,0.000061084924,0.00018322849,0.00044422157,0.008849259,0.0011166355],"genre_scores_gemma":[0.017957604,0.00012959709,0.97760314,0.000036865382,0.00002079782,0.00017200783,0.0009847186,0.0005049843,0.0025901848],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985405,0.00013025607,0.00015501428,0.0003850351,0.0006376172,0.0001515832],"domain_scores_gemma":[0.99604034,0.0017157173,0.00024256935,0.00076382124,0.0010862127,0.00015136336],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010283926,0.0025293808,0.0014851856,0.0061015915,0.001709963,0.0029117255,0.0029770234,0.0021131109,0.013326834],"category_scores_gemma":[0.0066440515,0.0012282565,0.0018817686,0.0056073745,0.0011806694,0.0055581657,0.0037811624,0.0024407313,0.0069838334],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004312409,0.00017363578,0.0010019735,0.00050376914,0.00008087197,0.00017036246,0.00027570085,0.012385171,0.025749896,0.032082118,0.015232778,0.91191256],"study_design_scores_gemma":[0.0003258117,0.00039720294,0.0012924786,0.0002573469,0.00021615908,0.0010423827,0.00049883476,0.6881991,0.087605365,0.15942068,0.060586575,0.00015805068],"about_ca_topic_score_codex":0.0045826305,"about_ca_topic_score_gemma":0.0062910016,"teacher_disagreement_score":0.013326834,"about_ca_system_score_codex":0.0012927381,"about_ca_system_score_gemma":0.0030955032,"threshold_uncertainty_score":0.044582665},"labels":[],"label_agreement":null},{"id":"W2125412731","doi":"10.1145/1639950.1640085","title":"Evaluation and usability of programming languages and tools (plateau)","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Usability; Programming language; Code refactoring; Program comprehension; Third-generation programming language; Software engineering; Fifth-generation programming language; Second-generation programming language; Fourth-generation programming language; Intersection (aeronautics); Programming paradigm; Inductive programming; Software; Human–computer interaction; Software system; Functional logic programming","score_opus":0.03441762922000367,"score_gpt":0.35143996243665326,"score_spread":0.3170223332166496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125412731","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.688836,0.015496162,0.20618407,0.0024429644,0.00052055845,0.006858586,0.0012797716,0.0018588629,0.07652309],"genre_scores_gemma":[0.9046418,0.0019123222,0.085903466,0.0003054523,0.000065945955,0.0029243268,0.00083193416,0.00031131404,0.003103335],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.77239656,0.14489284,0.01980215,0.0055517317,0.054931983,0.0024248515],"domain_scores_gemma":[0.62697345,0.2540803,0.019429622,0.0236211,0.07331713,0.0025784767],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10341441,0.0009127145,0.0010691584,0.008774433,0.0014810743,0.006889896,0.0013726876,0.0014404119,0.0020834664],"category_scores_gemma":[0.29497316,0.0004916626,0.0017836551,0.0062369453,0.0028357222,0.0063768113,0.0043165484,0.0011583174,0.00076619565],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020867551,0.00092817226,0.12865995,0.008195736,0.0012914987,0.00028040042,0.027799446,0.0031828564,0.007487338,0.01970762,0.007883499,0.79249674],"study_design_scores_gemma":[0.0012028554,0.015163189,0.5205286,0.017855559,0.0031713303,0.0032111127,0.061297566,0.05472099,0.049063604,0.08546608,0.18720464,0.0011144652],"about_ca_topic_score_codex":0.0022379626,"about_ca_topic_score_gemma":0.001621891,"teacher_disagreement_score":0.10341441,"about_ca_system_score_codex":0.002730354,"about_ca_system_score_gemma":0.0030930424,"threshold_uncertainty_score":0.5469142},"labels":[],"label_agreement":null},{"id":"W2125681674","doi":"10.1109/metric.2003.1232472","title":"An analogy-based approach for predicting design stability of Java classes","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Java; Testbed; Software; Software quality; Class (philosophy); Analogy; Compatibility (geochemistry); Stability (learning theory); Object-oriented programming; Software engineering; Software metric; Programming language; Artificial intelligence; Data mining; Machine learning; Software development; Engineering","score_opus":0.06407342464584245,"score_gpt":0.2978807540581682,"score_spread":0.23380732941232577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125681674","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2939697,0.00019496588,0.69757456,0.0005200345,0.000037703052,0.00031235936,0.0004394328,0.001963083,0.0049882815],"genre_scores_gemma":[0.86353534,0.000087116,0.13510576,0.000058669084,0.000014689417,0.00019046564,0.00030714218,0.000040981336,0.0006597796],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988527,0.00037572798,0.000072654555,0.00028614342,0.00035686378,0.000055863813],"domain_scores_gemma":[0.9922563,0.00540403,0.0007238198,0.0007363719,0.000746703,0.00013281386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016092935,0.0008380441,0.0005745397,0.0032480005,0.00055299996,0.00086524896,0.0015235918,0.0015520266,0.00240933],"category_scores_gemma":[0.023317974,0.00037097937,0.00069161004,0.001684734,0.0007072433,0.0020570024,0.0007282415,0.0008629543,0.00044978253],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049775885,0.0006819251,0.057428192,0.00025475287,0.00014061548,0.0005949207,0.00061075686,0.6658671,0.010496139,0.021644663,0.0016915418,0.24009162],"study_design_scores_gemma":[0.000014101161,0.00010219523,0.003149066,0.000008664925,0.000022249453,0.00009399666,0.000024816349,0.9858888,0.0011012494,0.009165528,0.00041525115,0.00001419408],"about_ca_topic_score_codex":0.0034469652,"about_ca_topic_score_gemma":0.0031839,"teacher_disagreement_score":0.0034469652,"about_ca_system_score_codex":0.001046622,"about_ca_system_score_gemma":0.0007240798,"threshold_uncertainty_score":0.008510888},"labels":[],"label_agreement":null},{"id":"W2125685892","doi":"10.1109/icsm.2004.1357822","title":"Evaluating similarity measures for software decompositions","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Dependency (UML); Cluster analysis; Similarity (geometry); Measure (data warehouse); Data mining; Dependency graph; Software; Similarity measure; Decomposition; Theoretical computer science; Graph; Artificial intelligence; Image (mathematics); Programming language","score_opus":0.0888671026905603,"score_gpt":0.3812133118091712,"score_spread":0.29234620911861087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125685892","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46880212,0.0036333776,0.52114356,0.00041153288,0.00017422183,0.00035931612,0.001170493,0.00071239914,0.003593017],"genre_scores_gemma":[0.7803577,0.00053520885,0.21594141,0.000049238788,0.000107577354,0.00025549534,0.0021101367,0.00012905762,0.00051421585],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9885867,0.0031208135,0.001472293,0.0016854492,0.0047480017,0.00038668976],"domain_scores_gemma":[0.96180135,0.024154263,0.004351726,0.0032390177,0.0051990515,0.0012545286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008728411,0.00081853976,0.0016050036,0.016350385,0.0009900253,0.00344597,0.0015357293,0.0017267608,0.0009811196],"category_scores_gemma":[0.06052173,0.00033875895,0.0009732465,0.008952188,0.0015539541,0.006118259,0.0025385525,0.0012268475,0.00034238579],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013903627,0.0004999776,0.092664205,0.0013546456,0.0010361684,0.00041771666,0.0020387506,0.12641591,0.019384168,0.06323237,0.006727771,0.6848379],"study_design_scores_gemma":[0.00011533913,0.0012363413,0.05050032,0.00020174019,0.0002066717,0.00087372836,0.0017224901,0.7664224,0.011429851,0.159074,0.008053586,0.0001635031],"about_ca_topic_score_codex":0.0013202941,"about_ca_topic_score_gemma":0.0013457162,"teacher_disagreement_score":0.016350385,"about_ca_system_score_codex":0.0018696534,"about_ca_system_score_gemma":0.0007007712,"threshold_uncertainty_score":0.046160817},"labels":[],"label_agreement":null},{"id":"W2125752271","doi":"10.5555/2664398.2664403","title":"Dispersion of changes in cloned and non-cloned code","year":2012,"lang":"en","type":"article","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Java; Computer science; Cloning (programming); Software maintenance; Asynchronous communication; Software evolution; Source lines of code; clone (Java method); Software; Programming language; Biology; Software system; Parallel computing; Genetics; Gene; Telecommunications","score_opus":0.04085499255652019,"score_gpt":0.3161615759757986,"score_spread":0.2753065834192784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125752271","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9775098,0.00025155966,0.019760735,0.000042272775,0.000029758814,0.000055298176,0.00044375358,0.0010937538,0.0008131062],"genre_scores_gemma":[0.9866196,0.00005634824,0.012009837,0.000013901699,0.000008882742,0.000034093664,0.00066376827,0.00012149903,0.00047196678],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9920507,0.001086239,0.00074599247,0.0014178331,0.0044078967,0.00029138167],"domain_scores_gemma":[0.9458583,0.024091993,0.0100075705,0.00788017,0.011208407,0.00095355476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031927289,0.00040073448,0.00040723328,0.0036159807,0.00040075352,0.0010835155,0.0007566913,0.00043215376,0.00061191333],"category_scores_gemma":[0.031227158,0.00033220593,0.000343986,0.0025946999,0.0005868602,0.0014537673,0.00085512374,0.0008357773,0.00018974075],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020542834,0.0006055035,0.43071443,0.00055312284,0.00045378797,0.00084974035,0.002885129,0.041013755,0.2263936,0.0021148615,0.0012672915,0.29109442],"study_design_scores_gemma":[0.000046533536,0.0014525848,0.46046844,0.00008587042,0.00022683556,0.001095357,0.0007352445,0.25361776,0.27599886,0.0020342476,0.0041097165,0.00012850405],"about_ca_topic_score_codex":0.0018375018,"about_ca_topic_score_gemma":0.0015965019,"teacher_disagreement_score":0.0036159807,"about_ca_system_score_codex":0.0006652305,"about_ca_system_score_gemma":0.00045636445,"threshold_uncertainty_score":0.016884923},"labels":[],"label_agreement":null},{"id":"W2125759561","doi":"10.1145/1368088.1368115","title":"On the difficulty of replicating human subjects studies in software engineering","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Replication (statistics); Comparability; Computer science; Context (archaeology); Literal (mathematical logic); Empirical research; Software; Test (biology); Artificial intelligence; Software engineering; Cognitive psychology; Programming language; Psychology; Mathematics; Statistics","score_opus":0.07825651788581889,"score_gpt":0.3125074568354242,"score_spread":0.2342509389496053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125759561","genre_codex":"methods","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10148197,0.0058471947,0.83412856,0.011170455,0.0074037947,0.022787118,0.0007211345,0.00097044295,0.0154894395],"genre_scores_gemma":[0.3814714,0.0011195113,0.5764634,0.008117389,0.0015312824,0.02871246,0.00045158944,0.00045315438,0.001679842],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.23960125,0.66677237,0.030680124,0.021344926,0.04012842,0.001472848],"domain_scores_gemma":[0.07145346,0.6813721,0.035031974,0.17658432,0.034578595,0.000979533],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.6270509,0.0030509052,0.003164387,0.003975099,0.005538221,0.0073789987,0.008972805,0.00813703,0.0032465346],"category_scores_gemma":[0.84221554,0.0022132327,0.0035847675,0.0042961575,0.01660033,0.014559018,0.006195382,0.008059822,0.0016731389],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012209256,0.00997311,0.1158165,0.018079605,0.015109382,0.007730332,0.09153187,0.021455398,0.02435245,0.2728906,0.027044224,0.3838073],"study_design_scores_gemma":[0.013832972,0.04054641,0.063863955,0.012803603,0.005374427,0.0051576,0.026348142,0.049716476,0.03804159,0.5677426,0.17514485,0.0014274146],"about_ca_topic_score_codex":0.00262286,"about_ca_topic_score_gemma":0.0022017718,"teacher_disagreement_score":0.37294912,"about_ca_system_score_codex":0.0032833433,"about_ca_system_score_gemma":0.0052932412,"threshold_uncertainty_score":0.45991272},"labels":[],"label_agreement":null},{"id":"W2125777058","doi":"10.1007/978-3-642-41641-5_2","title":"Visualizing and Measuring Enterprise Architecture: An Exploratory BioPharma Case","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Ottawa Mental Health Centre","funders":"","keywords":"Computer science; Enterprise architecture; Enterprise architecture framework; Business architecture; Enterprise architecture management; Architecture; Software engineering; Software; Software architecture; Engineering; Business process; Operating system; Operations management; Geography","score_opus":0.02530361081536024,"score_gpt":0.26013652865178755,"score_spread":0.23483291783642732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125777058","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98169875,0.00015107627,0.015249195,0.0002329046,0.0000063727275,0.0002706243,0.00043268452,0.00011887619,0.0018395196],"genre_scores_gemma":[0.92997265,0.00020470665,0.06685684,0.00007409086,0.000013538034,0.0002642276,0.0010571593,0.000079904996,0.0014768869],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99660873,0.0015694508,0.0002071291,0.00047810847,0.00087721756,0.00025941286],"domain_scores_gemma":[0.977772,0.017668484,0.001114284,0.0016609618,0.0013475228,0.00043671837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003033904,0.00069875724,0.000420572,0.004257117,0.001424281,0.0014152513,0.0012387371,0.0017018124,0.0012558822],"category_scores_gemma":[0.010822843,0.0004485695,0.0007612496,0.0034771718,0.0014737697,0.0018528803,0.0018355987,0.0010382952,0.00037425704],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016829426,0.007203835,0.4032172,0.0022002251,0.00037419435,0.0315475,0.10564778,0.0376939,0.09926133,0.015877612,0.0082272515,0.28706628],"study_design_scores_gemma":[0.0004085706,0.0034677566,0.45465326,0.0008480078,0.00041467784,0.018011145,0.07176959,0.2382479,0.14354922,0.02471354,0.043414395,0.00050199404],"about_ca_topic_score_codex":0.004387268,"about_ca_topic_score_gemma":0.0074795075,"teacher_disagreement_score":0.004387268,"about_ca_system_score_codex":0.001380544,"about_ca_system_score_gemma":0.00057608244,"threshold_uncertainty_score":0.016045034},"labels":[],"label_agreement":null},{"id":"W2125908974","doi":"10.1109/iceccs.2005.57","title":"Measuring Various Properties of Execution Traces to Help Build Better Trace Analysis Tools","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"TRACE (psycholinguistics); Computer science; Set (abstract data type); Software; Simple (philosophy); Software engineering; Programming language","score_opus":0.04640081999690251,"score_gpt":0.24942794322235354,"score_spread":0.20302712322545102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125908974","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2764207,0.0006584724,0.71371377,0.0005676879,0.00006153001,0.00065801095,0.0014071665,0.0034729135,0.0030398609],"genre_scores_gemma":[0.70479107,0.00040299247,0.29172713,0.000051995437,0.000035139456,0.00031146756,0.0018329112,0.00028945907,0.00055785436],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9948597,0.0018582691,0.0007268251,0.00052464573,0.0018172993,0.00021333616],"domain_scores_gemma":[0.945458,0.03212777,0.00674593,0.0064734872,0.00818986,0.0010049624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005210826,0.001507684,0.0010370852,0.0066071628,0.00061594835,0.0020530235,0.0010880044,0.0008907668,0.0018474972],"category_scores_gemma":[0.039519604,0.0003455688,0.0008105808,0.004149657,0.0007514886,0.004964388,0.00096451637,0.0013134654,0.00046139912],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070264266,0.0014726836,0.23155732,0.0018487381,0.0007573008,0.0005966471,0.0021640556,0.11001767,0.12186162,0.02409322,0.002605441,0.5023226],"study_design_scores_gemma":[0.00012178299,0.001858138,0.14709361,0.0005289886,0.00049983815,0.0011850348,0.0023132656,0.6571211,0.12931961,0.04686382,0.01276034,0.0003344116],"about_ca_topic_score_codex":0.002011747,"about_ca_topic_score_gemma":0.0026780292,"teacher_disagreement_score":0.0066071628,"about_ca_system_score_codex":0.0007801039,"about_ca_system_score_gemma":0.0013484361,"threshold_uncertainty_score":0.02755785},"labels":[],"label_agreement":null},{"id":"W2126036011","doi":"10.1145/2387358.2387360","title":"An empirical study on clone stability","year":2012,"lang":"en","type":"article","venue":"ACM SIGAPP Applied Computing Review","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Cloning (programming); clone (Java method); Java; Code (set theory); Computer science; Stability (learning theory); Source code; Programming language; Source lines of code; Software; Biology; Genetics; Set (abstract data type); Machine learning","score_opus":0.0823338992172196,"score_gpt":0.38100455216125284,"score_spread":0.2986706529440332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126036011","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.995323,0.0002652772,0.0017659006,0.00012614066,0.000004926094,0.00004667178,0.00015294712,0.00002323978,0.0022918528],"genre_scores_gemma":[0.99841034,0.000106924446,0.00083532825,0.000041387902,0.0000070374626,0.000037087775,0.00016485194,0.000012741412,0.00038428925],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98971623,0.0039936183,0.0009320875,0.0015239241,0.0034658278,0.00036824314],"domain_scores_gemma":[0.68368906,0.23576903,0.038099214,0.011810011,0.027901584,0.002731176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010289443,0.00025694544,0.00035086065,0.0034047945,0.000933316,0.0020149585,0.0009584558,0.0006863133,0.0025118145],"category_scores_gemma":[0.15067284,0.00023839051,0.0002477423,0.0036789547,0.0016311171,0.003508668,0.0015268321,0.00082994346,0.00037257324],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023732825,0.00019846015,0.93328124,0.00021847362,0.000070515765,0.0003774817,0.01567973,0.000688766,0.0013957238,0.0025553105,0.00074414787,0.044552796],"study_design_scores_gemma":[0.000024566883,0.0005830268,0.9564896,0.00015163455,0.00008068194,0.0010850343,0.018878063,0.008505441,0.0026594568,0.0023000585,0.009196162,0.000046350393],"about_ca_topic_score_codex":0.0021788876,"about_ca_topic_score_gemma":0.0016444103,"teacher_disagreement_score":0.010289443,"about_ca_system_score_codex":0.0012850435,"about_ca_system_score_gemma":0.0009579805,"threshold_uncertainty_score":0.054416418},"labels":[],"label_agreement":null},{"id":"W2126063715","doi":"10.1016/s0167-6423(99)00036-2","title":"How do program understanding tools affect how programmers understand programs?","year":2000,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":155,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; Simon Fraser University","funders":"","keywords":"Computer science; Program comprehension; Programmer; Source code; Software engineering; Human–computer interaction; Usability; Affect (linguistics); Comprehension; Software; Programming language; Data science; Software system","score_opus":0.06352619708920401,"score_gpt":0.29323639168980353,"score_spread":0.22971019460059952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126063715","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95623773,0.0004660661,0.021867402,0.0016870678,0.000055862012,0.00006088728,0.00010274235,0.0005999436,0.018922273],"genre_scores_gemma":[0.9919962,0.00014521503,0.005755283,0.00029169148,0.000014713773,0.000030185807,0.000072551484,0.0004116586,0.001282501],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99001384,0.0057791886,0.0003017991,0.0014080039,0.0018847456,0.0006123791],"domain_scores_gemma":[0.8258317,0.14062184,0.013393721,0.007109981,0.010460144,0.0025825445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006999307,0.00057153025,0.00045240662,0.0012302294,0.0009915588,0.0071937945,0.00091548206,0.002658373,0.003694713],"category_scores_gemma":[0.16852848,0.00083905447,0.0004537546,0.0009699252,0.00200332,0.013867848,0.0019335487,0.002641962,0.0007066537],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023696742,0.002404824,0.47798467,0.00092144264,0.00072223134,0.00080500904,0.09726113,0.004994939,0.075947784,0.031152153,0.009314342,0.29612187],"study_design_scores_gemma":[0.0003556443,0.0016120742,0.7501581,0.0005664442,0.0012551764,0.0010362457,0.055632528,0.03366573,0.051655345,0.07823149,0.025420893,0.0004104253],"about_ca_topic_score_codex":0.0019519364,"about_ca_topic_score_gemma":0.0025332575,"teacher_disagreement_score":0.0071937945,"about_ca_system_score_codex":0.00088258827,"about_ca_system_score_gemma":0.00086127856,"threshold_uncertainty_score":0.037016332},"labels":[],"label_agreement":null},{"id":"W2126198188","doi":"10.1109/wcre.2006.13","title":"An Orchestrated Multi-view Software Architecture Reconstruction Environment","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Reverse engineering; Software system; Software; Schema (genetic algorithms); Software engineering; Scope (computer science); Software design; Software architecture; Set (abstract data type); Software design description; Process (computing); Software development; Component-based software engineering; Artificial intelligence; Programming language; Machine learning","score_opus":0.014218264080247041,"score_gpt":0.2355345672066559,"score_spread":0.22131630312640885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126198188","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004127345,0.000029968567,0.98843735,0.00004060594,0.000006122307,0.00003627463,0.000029616402,0.00654097,0.00075176364],"genre_scores_gemma":[0.07878925,0.0000807007,0.91794,0.000051084226,0.000006362614,0.00009993262,0.00033384832,0.0009206501,0.0017780903],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987429,0.00038479688,0.00007831039,0.0002296886,0.00044140324,0.00012295693],"domain_scores_gemma":[0.99824905,0.0004444642,0.00011646739,0.00092756376,0.0001742879,0.00008814945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017131633,0.0007237273,0.0005351347,0.0009639348,0.00044779474,0.0016609519,0.0020475814,0.0010382095,0.0034956513],"category_scores_gemma":[0.0029000468,0.0009154368,0.0011972665,0.000548531,0.0008968312,0.002828028,0.0032358784,0.0017104921,0.001172116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006985762,0.000322225,0.004333942,0.0003560733,0.00015734049,0.0019235808,0.003270926,0.09798947,0.09540981,0.12539046,0.007940562,0.66220695],"study_design_scores_gemma":[0.00013647847,0.00028692334,0.0010092793,0.00009915251,0.00010000216,0.0016043724,0.0006622324,0.7659255,0.09912816,0.055163503,0.07575596,0.00012846153],"about_ca_topic_score_codex":0.000952048,"about_ca_topic_score_gemma":0.001375866,"teacher_disagreement_score":0.0034956513,"about_ca_system_score_codex":0.00040368084,"about_ca_system_score_gemma":0.0007907185,"threshold_uncertainty_score":0.011694133},"labels":[],"label_agreement":null},{"id":"W2126730858","doi":"10.1145/1806799.1806846","title":"Identifying crosscutting concerns using historical code changes","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Commit; Merge (version control); Computer science; Source code; Data science; Code review; Code (set theory); Complement (music); Open source; Software engineering; Static program analysis; Risk analysis (engineering); Programming language; Software; Software development; Database; Information retrieval; Business","score_opus":0.10264925049518206,"score_gpt":0.355441768671808,"score_spread":0.25279251817662596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126730858","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80490756,0.0014267623,0.17930478,0.00031281373,0.00009754541,0.00031945025,0.0025900851,0.005844021,0.0051969374],"genre_scores_gemma":[0.8502864,0.00077435654,0.140043,0.000062959094,0.000050309904,0.00013463988,0.005541875,0.000531112,0.0025753477],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9976776,0.00031301164,0.00018811897,0.00076286227,0.0009346165,0.00012379029],"domain_scores_gemma":[0.97907627,0.007978148,0.0042284112,0.0040991665,0.0041564764,0.00046148073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019678953,0.00071495544,0.0003803096,0.0055613923,0.000517286,0.0011285874,0.0010180876,0.00075364794,0.00093238195],"category_scores_gemma":[0.019291611,0.00049551873,0.00062868005,0.00282226,0.00038501408,0.0022883236,0.0009231292,0.0008548786,0.00046343644],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023493732,0.00023555337,0.31862134,0.00059092755,0.00017490624,0.0012831503,0.0027784435,0.009087708,0.03795897,0.0023764581,0.002942893,0.6237147],"study_design_scores_gemma":[0.000044415,0.0006887542,0.5954215,0.00028091282,0.0004722038,0.0040663155,0.0015800617,0.2698835,0.08014212,0.007244264,0.040005755,0.00017013462],"about_ca_topic_score_codex":0.004275465,"about_ca_topic_score_gemma":0.010566963,"teacher_disagreement_score":0.0055613923,"about_ca_system_score_codex":0.00053883134,"about_ca_system_score_gemma":0.0007385412,"threshold_uncertainty_score":0.010407329},"labels":[],"label_agreement":null},{"id":"W2126816472","doi":"10.1109/icsm.2006.30","title":"Guiding the Application of Design Patterns Based on UML Models","year":2006,"lang":"en","type":"article","venue":"Proceedings/Proceedings - Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Unified Modeling Language; Computer science; Software engineering; Applications of UML; Structural pattern; Software design pattern; Software design; Software; UML tool; Programming language; Software development","score_opus":0.04783375254324693,"score_gpt":0.25681082107045583,"score_spread":0.2089770685272089,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126816472","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0110395895,0.00012954314,0.9840625,0.00077385036,0.000030876872,0.00070274365,0.00007812277,0.001043853,0.002138959],"genre_scores_gemma":[0.033955388,0.00016135299,0.96383494,0.00012049416,0.000008590532,0.00050224725,0.00022703374,0.00022335324,0.00096664],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96658605,0.020393575,0.0030119256,0.0024521109,0.006483035,0.001073288],"domain_scores_gemma":[0.9465382,0.032012973,0.0056159305,0.007855757,0.0072052097,0.00077185873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026203394,0.0021934947,0.0010955401,0.004473268,0.0013562558,0.005689645,0.0035815174,0.002974444,0.0019943349],"category_scores_gemma":[0.08156997,0.0020180389,0.001382014,0.0024804932,0.0025989383,0.0062084636,0.0045755515,0.0029390263,0.0013904176],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003484756,0.0010118047,0.01444829,0.002562913,0.00025762356,0.0014890775,0.02512018,0.06582886,0.054579444,0.1636169,0.00830437,0.6624321],"study_design_scores_gemma":[0.00041290413,0.0005908986,0.0037570316,0.0026581045,0.0002542085,0.0018121703,0.007919799,0.57947826,0.06018184,0.19947876,0.14316218,0.00029380515],"about_ca_topic_score_codex":0.0045911954,"about_ca_topic_score_gemma":0.008559495,"teacher_disagreement_score":0.026203394,"about_ca_system_score_codex":0.0019613674,"about_ca_system_score_gemma":0.0060973708,"threshold_uncertainty_score":0.13857847},"labels":[],"label_agreement":null},{"id":"W2127045627","doi":"10.1109/wcre.2006.14","title":"Animated Visualization of Software History using Evolution Storyboards","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Storyboard; Computer science; Software evolution; Software visualization; Visualization; Software development; Software; Software system; Software engineering; Software construction; Programming language; Artificial intelligence; Multimedia","score_opus":0.022273858065523403,"score_gpt":0.2663185180922234,"score_spread":0.2440446600267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127045627","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11149296,0.00080615905,0.84463614,0.0008760786,0.00022005623,0.00024524078,0.005526858,0.023649814,0.012546602],"genre_scores_gemma":[0.46560138,0.0013508617,0.5204087,0.00014936534,0.00009646079,0.0003540309,0.0052511906,0.002951134,0.0038368609],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997043,0.000117183394,0.000023218901,0.000049655257,0.000081287995,0.000024267014],"domain_scores_gemma":[0.9976687,0.001427182,0.00020120502,0.00024363557,0.000292863,0.00016647966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006234887,0.0010325935,0.0003793582,0.003000696,0.00036267703,0.0017548526,0.00063750194,0.0006820523,0.008704391],"category_scores_gemma":[0.0040094135,0.0004688035,0.00056978513,0.001600725,0.0003318533,0.0019194859,0.001383267,0.0010055159,0.0006707831],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014434175,0.00034851063,0.010997096,0.001996308,0.00033047795,0.002460954,0.012723981,0.12125174,0.09354365,0.048431728,0.057175,0.64929706],"study_design_scores_gemma":[0.00034894934,0.00033118474,0.02347818,0.0006059431,0.00020277833,0.00110694,0.0017582555,0.6660851,0.047161218,0.05110871,0.20749803,0.0003147404],"about_ca_topic_score_codex":0.0019212507,"about_ca_topic_score_gemma":0.0019804828,"teacher_disagreement_score":0.008704391,"about_ca_system_score_codex":0.0003531161,"about_ca_system_score_gemma":0.0003370654,"threshold_uncertainty_score":0.029119074},"labels":[],"label_agreement":null},{"id":"W2127176799","doi":"10.1109/icac.2005.49","title":"Quickly Finding Known Software Problems via Automated Symptom Matching","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":81,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Computer science; Matching (statistics); Software; Component (thermodynamics); Pattern matching; Software system; Product (mathematics); Software product line; Data mining; Artificial intelligence; Programming language; Software development; Mathematics","score_opus":0.015085838354808622,"score_gpt":0.26699250249911727,"score_spread":0.2519066641443086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127176799","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10459504,0.00044954766,0.8413517,0.00055431976,0.00007198685,0.0004934425,0.00091352104,0.049280696,0.0022897564],"genre_scores_gemma":[0.26622355,0.00019686588,0.7289425,0.00017994037,0.000036255944,0.00019071109,0.001949621,0.0006099676,0.001670556],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968732,0.00062582194,0.00033997386,0.00079112925,0.0012045547,0.00016532117],"domain_scores_gemma":[0.9899455,0.004789161,0.0016698329,0.0019237985,0.0013665981,0.0003051227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020514415,0.0011372253,0.0014961198,0.004485492,0.000877062,0.0018713644,0.0033078163,0.0018753624,0.0044020186],"category_scores_gemma":[0.013532861,0.0006507625,0.0008972441,0.0025241252,0.0006983352,0.0037738013,0.0026652466,0.0012156027,0.0018293111],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067351985,0.0008939485,0.022666987,0.0006891405,0.00021434395,0.00093440776,0.0010709676,0.013042342,0.067960724,0.0049219197,0.012643628,0.8742882],"study_design_scores_gemma":[0.00034947685,0.00091418973,0.018034352,0.00011742565,0.00027428346,0.0034627093,0.0011414008,0.80120283,0.12527731,0.027452117,0.021563545,0.00021042852],"about_ca_topic_score_codex":0.0030020017,"about_ca_topic_score_gemma":0.003530296,"teacher_disagreement_score":0.004485492,"about_ca_system_score_codex":0.0005456515,"about_ca_system_score_gemma":0.0013470587,"threshold_uncertainty_score":0.014726222},"labels":[],"label_agreement":null},{"id":"W2127190390","doi":"10.1145/1806799.1806856","title":"A degree-of-knowledge model to capture source code familiarity","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":121,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; KPI-driven code analysis; Source code; Code (set theory); Code review; Software; Software development; Software quality; Value (mathematics); Software engineering; Programming language; Machine learning","score_opus":0.041004224503487896,"score_gpt":0.29336407727349645,"score_spread":0.25235985277000855,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127190390","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1443786,0.00033699835,0.8433385,0.00070331735,0.00004060739,0.00023485409,0.0010544771,0.00064024644,0.00927245],"genre_scores_gemma":[0.8982233,0.00012771961,0.09856256,0.00007740273,0.000035730445,0.00020934343,0.0006783142,0.00006762926,0.002018014],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955903,0.0013238846,0.0003383628,0.0012826292,0.0010767202,0.00038807676],"domain_scores_gemma":[0.95220804,0.035504498,0.0036887198,0.0046854075,0.002963927,0.00094951375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004485081,0.0007501792,0.00080514286,0.0054008695,0.00080353254,0.003086769,0.0017213393,0.0019021155,0.0036297215],"category_scores_gemma":[0.047291405,0.00050365715,0.0016688498,0.004966784,0.0018303362,0.010965803,0.0020244622,0.002041344,0.0010451857],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082405325,0.0011294235,0.20489787,0.00063430425,0.0006281647,0.000522086,0.0048650336,0.2600947,0.007315507,0.21455805,0.0061099934,0.2984209],"study_design_scores_gemma":[0.000050664945,0.00021381697,0.034689657,0.000064001855,0.00012831105,0.0007691757,0.00044720998,0.79916406,0.0021886588,0.15675229,0.005425796,0.000106292166],"about_ca_topic_score_codex":0.008294485,"about_ca_topic_score_gemma":0.00773433,"teacher_disagreement_score":0.008294485,"about_ca_system_score_codex":0.0026443389,"about_ca_system_score_gemma":0.0013451165,"threshold_uncertainty_score":0.023719668},"labels":[],"label_agreement":null},{"id":"W2127424531","doi":"","title":"A system of references for software measurements with ISO 19761 (COSMIC-FFP)","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software measurement; Software; Metrology; Computer science; Verification and validation; Reliability engineering; Repeatability; Software engineering; Systems engineering; Software quality; Software development; Engineering; Mathematics; Programming language; Statistics","score_opus":0.06056794410348576,"score_gpt":0.27415349940686473,"score_spread":0.21358555530337897,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127424531","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036937124,0.0014585548,0.9421968,0.0008218122,0.001108916,0.0014592415,0.0026117088,0.012822467,0.033826753],"genre_scores_gemma":[0.020184126,0.0009925986,0.9543204,0.00045384796,0.00029908706,0.0025993155,0.008617258,0.0020076635,0.010525737],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9705437,0.0071715885,0.0049867094,0.003082427,0.013339783,0.00087584555],"domain_scores_gemma":[0.9688928,0.0035125404,0.002400548,0.0073817237,0.017220758,0.0005916698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015996518,0.0022033856,0.0014914264,0.014742266,0.0036638982,0.00679881,0.004557312,0.0046412623,0.012429561],"category_scores_gemma":[0.041934278,0.0010578266,0.0017541677,0.011473561,0.0026996492,0.0057593454,0.0048589786,0.0033768788,0.0144251855],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026146293,0.00027689728,0.0041700075,0.0014193968,0.000077938836,0.00052296004,0.0023697664,0.0054859472,0.01505747,0.2526982,0.11648023,0.60117966],"study_design_scores_gemma":[0.0000592473,0.0003042192,0.002230472,0.0009381798,0.00009880443,0.0006744888,0.00043109496,0.009286895,0.013553826,0.03073348,0.94151014,0.00017911536],"about_ca_topic_score_codex":0.007975701,"about_ca_topic_score_gemma":0.005882092,"teacher_disagreement_score":0.015996518,"about_ca_system_score_codex":0.0037816328,"about_ca_system_score_gemma":0.011614861,"threshold_uncertainty_score":0.08459866},"labels":[],"label_agreement":null},{"id":"W2127561942","doi":"10.1109/wcre.2013.6671325","title":"An approach to clone detection in behavioural models","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Computer science; Artificial intelligence; Biology; Genetics; Gene","score_opus":0.03998833481447143,"score_gpt":0.2617778283901415,"score_spread":0.2217894935756701,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127561942","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0028505633,0.000044289147,0.99502844,0.00011212525,0.0000160172,0.000068557005,0.00005569049,0.0013069755,0.00051745144],"genre_scores_gemma":[0.076273106,0.000117038755,0.91923153,0.00022861442,0.000037811114,0.00026975077,0.000444318,0.0007514022,0.0026464725],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9889614,0.0029409162,0.00077800656,0.0019795736,0.004870882,0.00046928728],"domain_scores_gemma":[0.97254974,0.011034302,0.0025500976,0.008500074,0.004949017,0.00041673728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052523525,0.0012490519,0.0010522762,0.0035476312,0.0013997343,0.0032368507,0.0033537883,0.002930253,0.0022726636],"category_scores_gemma":[0.037338316,0.0014419244,0.0027998681,0.0018051184,0.0025346207,0.005714154,0.0058900076,0.004318752,0.0010503173],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039335398,0.0003241535,0.016463568,0.00092837034,0.00028341255,0.0019068952,0.010693637,0.056023557,0.06486782,0.38207936,0.006146667,0.45988914],"study_design_scores_gemma":[0.00005209978,0.00024024458,0.0026428506,0.0003409605,0.00021408242,0.0019908925,0.0012512512,0.6007445,0.05768528,0.2561914,0.07844789,0.00019855851],"about_ca_topic_score_codex":0.004037023,"about_ca_topic_score_gemma":0.004712982,"teacher_disagreement_score":0.0052523525,"about_ca_system_score_codex":0.0018615633,"about_ca_system_score_gemma":0.0023511997,"threshold_uncertainty_score":0.027777374},"labels":[],"label_agreement":null},{"id":"W2127619534","doi":"10.1109/wcre.1998.723173","title":"Evaluating architectural extractors","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Workbench; Computer science; Software engineering; Reverse engineering; Software; Visualization; Source code; Architectural pattern; Software system; Software construction; Programming language; Artificial intelligence","score_opus":0.09881032934449571,"score_gpt":0.344662248214615,"score_spread":0.24585191887011928,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127619534","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39805084,0.009977566,0.5070906,0.00097467616,0.0003395264,0.0021112869,0.011535193,0.046274204,0.023646027],"genre_scores_gemma":[0.30064997,0.0034523197,0.6533408,0.00018355488,0.00009929361,0.0006419241,0.030324424,0.0026793696,0.0086284],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98808473,0.00290268,0.0012860037,0.0012389149,0.0060425145,0.00044510793],"domain_scores_gemma":[0.9210379,0.055000145,0.004048837,0.008548774,0.01071358,0.00065071654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009925081,0.0022920193,0.0015032493,0.009859324,0.0008735499,0.0037851322,0.0022360065,0.0018946507,0.00629373],"category_scores_gemma":[0.062941246,0.0007322235,0.001919326,0.005490589,0.0009239013,0.006749659,0.0025712412,0.0013089366,0.0033120164],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011111823,0.00044512723,0.017233964,0.002405524,0.00031657945,0.00043518635,0.0011119456,0.019771867,0.016879318,0.0042498326,0.008296316,0.92774314],"study_design_scores_gemma":[0.0014575064,0.0066340105,0.048926096,0.0017418102,0.0027727978,0.004134352,0.006998299,0.47588918,0.26932397,0.021662284,0.15991235,0.00054740947],"about_ca_topic_score_codex":0.0024429986,"about_ca_topic_score_gemma":0.0046759974,"teacher_disagreement_score":0.009925081,"about_ca_system_score_codex":0.0012204614,"about_ca_system_score_gemma":0.0019374166,"threshold_uncertainty_score":0.05248946},"labels":[],"label_agreement":null},{"id":"W2127721446","doi":"10.1109/ccece.2003.1226144","title":"A Web-based software engineering measurement expert system","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Software engineering; Software system; Software measurement; Expert system; Software construction; Verification and validation; Inference engine; Software development; Software; Data mining; Systems engineering; Engineering; Artificial intelligence; Operating system","score_opus":0.024429332386184763,"score_gpt":0.23173502825434936,"score_spread":0.2073056958681646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127721446","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01597868,0.00020173966,0.8357601,0.00034911552,0.000106826636,0.00076716725,0.0015073586,0.12859264,0.01673636],"genre_scores_gemma":[0.18970832,0.00028398843,0.77576405,0.0008090031,0.000115516414,0.0012975603,0.0063553792,0.0030424553,0.02262378],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99744797,0.00046685032,0.00028715414,0.00051754236,0.0011887212,0.00009185178],"domain_scores_gemma":[0.9927209,0.0021185928,0.0004410937,0.0012792569,0.0030271842,0.00041303053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027919766,0.0006037482,0.0010261822,0.0020134682,0.0003801403,0.0015718754,0.001765524,0.0011552761,0.011755232],"category_scores_gemma":[0.01114733,0.0005303253,0.0003909667,0.001226049,0.00028015435,0.0019981975,0.0014038867,0.0012505548,0.0074563264],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006470033,0.0009122991,0.0042027673,0.00047296353,0.00016620269,0.0005621734,0.00030517302,0.019498695,0.03637031,0.007980556,0.046146225,0.8827356],"study_design_scores_gemma":[0.00059920526,0.0005278092,0.011106669,0.00026580066,0.0002780598,0.0017609135,0.00015496978,0.71662676,0.059105564,0.023078768,0.18620132,0.00029419485],"about_ca_topic_score_codex":0.0012639674,"about_ca_topic_score_gemma":0.0012369627,"teacher_disagreement_score":0.011755232,"about_ca_system_score_codex":0.000535314,"about_ca_system_score_gemma":0.0014522148,"threshold_uncertainty_score":0.039325178},"labels":[],"label_agreement":null},{"id":"W2127847952","doi":"10.1109/fosm.2008.4659257","title":"Remixing visualization to support collaboration in software maintenance","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Visualization; Computer science; Software visualization; Software; Software engineering; Software development; Software analytics; Software maintenance; Process (computing); Collaborative software; Human–computer interaction; Software construction; World Wide Web; Data mining","score_opus":0.020858944337745572,"score_gpt":0.30531924209016936,"score_spread":0.28446029775242376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127847952","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040068857,0.0033336733,0.93287057,0.0025682307,0.00032128353,0.000183907,0.000118521726,0.0073237303,0.013211261],"genre_scores_gemma":[0.2435627,0.0020499483,0.7489453,0.0002576711,0.00023202135,0.00019958694,0.00023961627,0.0006388332,0.0038743848],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979436,0.0010809678,0.00014036919,0.00026714397,0.0004520322,0.00011587656],"domain_scores_gemma":[0.991508,0.0040850877,0.00094434014,0.0019993235,0.0011303306,0.00033291217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039503984,0.0007940672,0.00061843015,0.0021558641,0.0009419831,0.0027656395,0.0019820207,0.0012852357,0.003969194],"category_scores_gemma":[0.013751424,0.00051830796,0.0007221572,0.0017578271,0.0008855125,0.0047047427,0.0024013517,0.001546762,0.0009415092],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029666396,0.00029085326,0.0058553,0.0011503119,0.00014609881,0.00089246733,0.010742039,0.022418825,0.039855737,0.1075848,0.02069704,0.79006994],"study_design_scores_gemma":[0.0003435521,0.0012156434,0.010699009,0.0012605719,0.0005580457,0.003516524,0.0026556181,0.33197185,0.076602064,0.20943594,0.3613603,0.00038087508],"about_ca_topic_score_codex":0.0012732004,"about_ca_topic_score_gemma":0.0015523916,"teacher_disagreement_score":0.003969194,"about_ca_system_score_codex":0.0004952271,"about_ca_system_score_gemma":0.0007270939,"threshold_uncertainty_score":0.020891964},"labels":[],"label_agreement":null},{"id":"W2127896511","doi":"10.1109/csmr.2002.995810","title":"On the role of design patterns in quality-driven re-engineering","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Restructuring; Computer science; Structural pattern; Software engineering; Software design pattern; Model-driven architecture; Engineering design process; Quality (philosophy); Perspective (graphical); Software; Software design; Systems engineering; Software development; Engineering; Artificial intelligence; Programming language","score_opus":0.03840666140903647,"score_gpt":0.278248749075265,"score_spread":0.2398420876662285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127896511","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062100567,0.0022275406,0.901185,0.00932902,0.00014263268,0.00024311764,0.000073272866,0.0004906306,0.024208289],"genre_scores_gemma":[0.4765915,0.0018896173,0.5153515,0.00079526915,0.000086261585,0.00038308825,0.00011112529,0.00024608627,0.004545601],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98200536,0.009335118,0.0014203324,0.001484151,0.0049229204,0.00083217677],"domain_scores_gemma":[0.9312637,0.043406673,0.006249346,0.011393367,0.0065531447,0.0011338243],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023600161,0.0008709283,0.00072426064,0.0027700514,0.002060093,0.0064537935,0.0017675507,0.0025778215,0.002225169],"category_scores_gemma":[0.06966165,0.0011940701,0.00091487117,0.0039235633,0.008831363,0.014668443,0.0033336622,0.0029479603,0.0004971443],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000092998096,0.0001093079,0.007592347,0.000402562,0.000056266937,0.00030504566,0.0060049873,0.018176889,0.001753924,0.7785363,0.0020880688,0.1848813],"study_design_scores_gemma":[0.000104045816,0.00024108902,0.002529283,0.0004990614,0.00010856116,0.0008413173,0.0027486128,0.0693931,0.0030169184,0.8742308,0.046191882,0.00009534205],"about_ca_topic_score_codex":0.004052996,"about_ca_topic_score_gemma":0.0039972053,"teacher_disagreement_score":0.023600161,"about_ca_system_score_codex":0.0027934283,"about_ca_system_score_gemma":0.0028798897,"threshold_uncertainty_score":0.12481105},"labels":[],"label_agreement":null},{"id":"W2127916964","doi":"10.1145/2597073.2597104","title":"Prediction and ranking of co-change candidates for clones","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Fragment (logic); Programmer; Computer science; Code (set theory); Group (periodic table); Programming language; Biology; Computational biology; Genetics; DNA; Chemistry","score_opus":0.0358089874994088,"score_gpt":0.27850423457236106,"score_spread":0.24269524707295226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127916964","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9280937,0.003335457,0.05762829,0.00047190677,0.00019119488,0.0003454925,0.0034036993,0.0036739958,0.0028563174],"genre_scores_gemma":[0.90537214,0.0007252871,0.075950935,0.00018110729,0.00014924639,0.00023455628,0.012210593,0.000479342,0.0046967193],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99644655,0.0003712807,0.00029714085,0.0010947215,0.0013959915,0.00039440463],"domain_scores_gemma":[0.9822828,0.0074121417,0.0023437636,0.0011647559,0.0051668165,0.0016296995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019296831,0.0013726464,0.0011786731,0.010774291,0.0012811206,0.0019444062,0.0017651074,0.0021310374,0.002797871],"category_scores_gemma":[0.0133329425,0.0004355722,0.0017568668,0.0044904836,0.0007580161,0.001904689,0.001071427,0.0009857574,0.0018000219],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015391922,0.0006149243,0.65794355,0.0007032662,0.00044914742,0.0025410505,0.000639211,0.020088533,0.03261477,0.0023119114,0.017930258,0.26262417],"study_design_scores_gemma":[0.00026471788,0.0010753054,0.2632792,0.00015817115,0.0010170523,0.005420634,0.001195191,0.66128826,0.04390367,0.004075467,0.018144893,0.00017747472],"about_ca_topic_score_codex":0.008357616,"about_ca_topic_score_gemma":0.012420152,"teacher_disagreement_score":0.010774291,"about_ca_system_score_codex":0.00095221447,"about_ca_system_score_gemma":0.0016821241,"threshold_uncertainty_score":0.016617954},"labels":[],"label_agreement":null},{"id":"W2128138035","doi":"10.1109/wpc.2002.1021340","title":"Fused data-centric visualizations for software evolution environments","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Visualization; Abstraction; Software visualization; Hierarchy; Suite; Software evolution; Software; Software engineering; Human–computer interaction; Data visualization; Software system; Software architecture; Data science; Component-based software engineering; Software construction; Data mining; Programming language","score_opus":0.039981550144403924,"score_gpt":0.3007496325417957,"score_spread":0.2607680823973918,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128138035","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012407905,0.00030804545,0.9773588,0.00039311335,0.00006304419,0.000057512636,0.0003685652,0.0071758735,0.0018671058],"genre_scores_gemma":[0.1263052,0.00038013654,0.87044144,0.00007111818,0.000028409815,0.00012941034,0.0007584963,0.000986334,0.0008994687],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985322,0.0007128477,0.00010085032,0.0001392191,0.00043526987,0.00007961848],"domain_scores_gemma":[0.9951283,0.0024477046,0.0002896691,0.00092879706,0.00088684395,0.00031861596],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024096498,0.001209626,0.0006286474,0.0031320758,0.00085872307,0.0036505018,0.0010929967,0.0011173941,0.0052587716],"category_scores_gemma":[0.012769282,0.0007061397,0.0008777647,0.0022652852,0.0007063524,0.004619091,0.0038419983,0.0018120144,0.000932421],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011749669,0.00025084792,0.005655787,0.0010654242,0.00018266204,0.0010578057,0.008500494,0.10217697,0.047325328,0.21988215,0.030503629,0.58222395],"study_design_scores_gemma":[0.00019473079,0.00023436076,0.0025285003,0.0003041813,0.000116043986,0.0009851546,0.0009657177,0.6586613,0.030915022,0.16299425,0.14192273,0.00017799987],"about_ca_topic_score_codex":0.0020350532,"about_ca_topic_score_gemma":0.0025697812,"teacher_disagreement_score":0.0052587716,"about_ca_system_score_codex":0.00078303984,"about_ca_system_score_gemma":0.0007716836,"threshold_uncertainty_score":0.01759231},"labels":[],"label_agreement":null},{"id":"W2128152347","doi":"10.1109/cmpsac.2001.960601","title":"MOOSE - a task-driven program comprehension environment","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Concordia University","funders":"","keywords":"Program comprehension; Wizard; Computer science; Software visualization; Source code; Program slicing; Comprehension; Visualization; Task (project management); Software engineering; Human–computer interaction; Reverse engineering; Software; Programming language; Software system; World Wide Web; Software construction; Artificial intelligence; Systems engineering; Engineering","score_opus":0.026387752981410866,"score_gpt":0.24374215311068184,"score_spread":0.21735440012927099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128152347","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0145403845,0.000103525905,0.9074028,0.00020221621,0.000024966213,0.0003758519,0.0010647593,0.073407315,0.0028781877],"genre_scores_gemma":[0.084691755,0.0003400654,0.89767736,0.0002440847,0.000036460282,0.0010776339,0.0034733892,0.00666065,0.0057987203],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.999236,0.00023859077,0.000065788845,0.00018026118,0.00020357793,0.0000758131],"domain_scores_gemma":[0.9947936,0.0038072173,0.00028936297,0.00049641955,0.00040176578,0.00021161647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018497814,0.0013507226,0.00072676834,0.0009603524,0.00037900553,0.0018111728,0.002234122,0.0014045715,0.010575607],"category_scores_gemma":[0.009659035,0.0008258951,0.0011832719,0.0004644908,0.00064646977,0.0031281835,0.0024053277,0.0018227233,0.002831039],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002673496,0.0012278997,0.00724254,0.0020321894,0.00034528272,0.0019175224,0.007849248,0.0620987,0.11053825,0.066074006,0.08175524,0.65624565],"study_design_scores_gemma":[0.0010669328,0.0011042596,0.0059076785,0.0004997174,0.00018620037,0.001450922,0.00082880596,0.47327328,0.09106651,0.124855004,0.29931712,0.0004435777],"about_ca_topic_score_codex":0.0011599349,"about_ca_topic_score_gemma":0.001691618,"teacher_disagreement_score":0.010575607,"about_ca_system_score_codex":0.00030953268,"about_ca_system_score_gemma":0.0010130632,"threshold_uncertainty_score":0.035378993},"labels":[],"label_agreement":null},{"id":"W2128223360","doi":"10.1109/icpc.2006.40","title":"Programmer-friendly Decompiled Java","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Programming language; Java; AspectJ; Compiler; Java Modeling Language; Class (philosophy); Java annotation; Programmer; Generics in Java; Real time Java; Software; Aspect-oriented programming; Artificial intelligence","score_opus":0.008768505525653185,"score_gpt":0.2530855059548508,"score_spread":0.2443170004291976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128223360","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024106283,0.00031979886,0.7279117,0.0004529961,0.0003520912,0.00040964907,0.0016127193,0.22640753,0.018427223],"genre_scores_gemma":[0.13401179,0.00045803864,0.7329856,0.00088198786,0.00010543682,0.00038284573,0.0062579922,0.09369438,0.03122188],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978136,0.00022617048,0.00019629148,0.00048858975,0.0010369469,0.00023835726],"domain_scores_gemma":[0.99142075,0.0021236294,0.00046518666,0.0036293908,0.002028813,0.00033228466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018068425,0.001451348,0.00072493695,0.0011780945,0.000605449,0.0020395506,0.0019382825,0.00089461956,0.010140104],"category_scores_gemma":[0.008687017,0.001216841,0.0011107356,0.0005743122,0.00095984974,0.0028717893,0.0034549139,0.004180421,0.009836207],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007586642,0.0005891097,0.005784345,0.0011880465,0.000117113275,0.001829369,0.0019425,0.015358602,0.2213639,0.04646649,0.100247234,0.6043546],"study_design_scores_gemma":[0.00021341683,0.0002079974,0.0035178436,0.00040707158,0.00008150564,0.002321998,0.0002726474,0.12122843,0.39638066,0.03947111,0.4356281,0.00026926314],"about_ca_topic_score_codex":0.0007920857,"about_ca_topic_score_gemma":0.0015380863,"teacher_disagreement_score":0.010140104,"about_ca_system_score_codex":0.00070029753,"about_ca_system_score_gemma":0.0013520946,"threshold_uncertainty_score":0.033921957},"labels":[],"label_agreement":null},{"id":"W2128231994","doi":"10.1109/csmr.2004.1281409","title":"Integrating a reverse engineering tool with microsoft visual studio .NET","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Microsoft Visual Studio; Computer science; Interoperability; Reverse engineering; .NET Framework; Software engineering; Component (thermodynamics); Embedding; Software; Net (polyhedron); Systems engineering; Operating system; Engineering; Artificial intelligence","score_opus":0.007419439446489198,"score_gpt":0.2416095227428396,"score_spread":0.23419008329635038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128231994","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002337886,0.00015629026,0.9404953,0.00019968126,0.00017170972,0.00017869567,0.00019784638,0.049250644,0.007011822],"genre_scores_gemma":[0.017852357,0.00034073254,0.96371,0.00016523342,0.000043594508,0.00031475245,0.0007138438,0.008627623,0.00823183],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9973598,0.00053660007,0.00035889557,0.00034395242,0.0012354903,0.0001652806],"domain_scores_gemma":[0.99295276,0.0031584976,0.00058309385,0.0017471267,0.0013436205,0.00021494276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051105884,0.0010887085,0.00061254087,0.0021322942,0.00041030877,0.0025937855,0.0022922428,0.001163709,0.011482082],"category_scores_gemma":[0.011384614,0.0011590691,0.001041514,0.0011126922,0.0005472419,0.0029134236,0.0018880022,0.0028193854,0.0085352305],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047183858,0.00040481263,0.0014181755,0.0013316302,0.00016495811,0.0009172598,0.0008458192,0.0039342665,0.06301788,0.034363925,0.034006428,0.85912305],"study_design_scores_gemma":[0.00046320885,0.00054014125,0.0021043066,0.00067281217,0.0002767308,0.0043401844,0.00035319926,0.10712006,0.19453777,0.0503189,0.638969,0.00030365668],"about_ca_topic_score_codex":0.000644327,"about_ca_topic_score_gemma":0.0009921805,"teacher_disagreement_score":0.011482082,"about_ca_system_score_codex":0.00038989086,"about_ca_system_score_gemma":0.001392599,"threshold_uncertainty_score":0.03841144},"labels":[],"label_agreement":null},{"id":"W2128679831","doi":"10.1109/wcre.2001.957810","title":"Union schemas as a basis for a C++ extractor","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Queen's University","funders":"","keywords":"Computer science; Extractor; Programming language; Schema (genetic algorithms); Database schema; Schema migration; Software; Schema matching; Database; Theoretical computer science; Information retrieval; Semi-structured model; Database design","score_opus":0.04015817791239411,"score_gpt":0.28154446512548054,"score_spread":0.24138628721308641,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128679831","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002403237,0.00012899526,0.9670196,0.00018262638,0.00008031198,0.00018829206,0.00091283594,0.021917619,0.007166445],"genre_scores_gemma":[0.027274068,0.00027962786,0.9522746,0.0004280003,0.000079499965,0.0003220038,0.0033906759,0.009726979,0.00622447],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9950322,0.0007064337,0.0008491069,0.0008904472,0.002239218,0.00028259613],"domain_scores_gemma":[0.99186546,0.002086375,0.00050086796,0.0029289678,0.002429997,0.0001883535],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004983932,0.0012471059,0.0010715541,0.0027380201,0.0014453139,0.0053615565,0.0029003571,0.0013676683,0.016159376],"category_scores_gemma":[0.0118962815,0.0020360306,0.0024368775,0.0031913908,0.0023217616,0.008911173,0.004240238,0.00414051,0.009184208],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005232077,0.0001311331,0.002441929,0.0006777824,0.00016322409,0.0006438419,0.0012642919,0.012100311,0.019326167,0.67808217,0.04765426,0.23699173],"study_design_scores_gemma":[0.0001187339,0.00014945168,0.00073032203,0.00031351444,0.00017116459,0.0013166951,0.00032934258,0.06905002,0.1396496,0.13659053,0.65140104,0.00017958094],"about_ca_topic_score_codex":0.003166891,"about_ca_topic_score_gemma":0.0024633878,"teacher_disagreement_score":0.016159376,"about_ca_system_score_codex":0.0015150568,"about_ca_system_score_gemma":0.0031370074,"threshold_uncertainty_score":0.05405855},"labels":[],"label_agreement":null},{"id":"W2128683249","doi":"10.1109/promise.2007.5","title":"Decision Support Analysis for Software Effort Estimation by Analogy","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Analogy; Selection (genetic algorithm); Weighting; Context (archaeology); Process (computing); Adaptation (eye); Decision support system; Similarity (geometry); Decision analysis; Machine learning; Software; Artificial intelligence; Personalization; Data mining; Mathematics","score_opus":0.014357967037145349,"score_gpt":0.3080868266492115,"score_spread":0.29372885961206613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128683249","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034485035,0.00029187562,0.9590568,0.00044302488,0.00003210581,0.00032513848,0.00012148124,0.0002440967,0.005000413],"genre_scores_gemma":[0.46078208,0.00033359375,0.5365775,0.0001022831,0.000061687046,0.000817507,0.00025088247,0.000050736086,0.001023742],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98283017,0.011534595,0.00073451805,0.0008823766,0.0036469249,0.00037141205],"domain_scores_gemma":[0.9385224,0.05442516,0.0020715287,0.0013842557,0.0032495796,0.000346978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012845118,0.0016840285,0.001850145,0.004863123,0.0010771742,0.0030432525,0.0014884307,0.001642044,0.0077726557],"category_scores_gemma":[0.06858286,0.00049399474,0.0014854925,0.0037250374,0.0012888722,0.0036304735,0.0019710718,0.0018872584,0.0006676863],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063266733,0.0005481515,0.004674166,0.00081891974,0.0002671174,0.0004076182,0.00077579665,0.39616776,0.0020683056,0.31324777,0.002333031,0.27805868],"study_design_scores_gemma":[0.000072790404,0.00018891381,0.00058533624,0.00006960323,0.00005342093,0.00006990178,0.00014274179,0.8934166,0.00097645447,0.10264696,0.0017403406,0.000036991423],"about_ca_topic_score_codex":0.0018821992,"about_ca_topic_score_gemma":0.001065156,"teacher_disagreement_score":0.012845118,"about_ca_system_score_codex":0.0021827226,"about_ca_system_score_gemma":0.0020386463,"threshold_uncertainty_score":0.06793225},"labels":[],"label_agreement":null},{"id":"W2128702459","doi":"10.1142/s0219525914500064","title":"RECODE: SOFTWARE PACKAGE REFACTORING VIA COMMUNITY DETECTION IN BIPARTITE SOFTWARE NETWORKS","year":2014,"lang":"en","type":"article","venue":"Advances in Complex Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Key Research and Development Program of China; McMaster University","keywords":"Code refactoring; Computer science; Bipartite graph; Software; Set (abstract data type); Software package; Software maintenance; Software metric; Software engineering; Data mining; Software system; Software construction; Theoretical computer science; Programming language; Graph","score_opus":0.031390225217765556,"score_gpt":0.2889413635581913,"score_spread":0.2575511383404257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128702459","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029982146,0.00024435893,0.9665613,0.0001644764,0.000030219066,0.00013672409,0.000174022,0.001702334,0.0010044402],"genre_scores_gemma":[0.32035437,0.00023558328,0.6742145,0.00018247902,0.00005645609,0.00023081903,0.0012159005,0.00026182286,0.0032481034],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980854,0.00068027776,0.00006320266,0.00047363894,0.00053105335,0.00016649457],"domain_scores_gemma":[0.9962476,0.0017101275,0.00051826367,0.00050066784,0.0008502357,0.0001731627],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002084959,0.0010483658,0.0009311563,0.0043887952,0.00094836374,0.000965094,0.0019921076,0.001439376,0.0011925214],"category_scores_gemma":[0.0075006355,0.0004971111,0.0011040963,0.0021884951,0.00077093096,0.0018609947,0.0018399566,0.0009780052,0.00055],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003752374,0.00037697668,0.011791301,0.0004555369,0.00023930533,0.00043026268,0.0005458683,0.16915898,0.023537103,0.017982163,0.008161701,0.7669455],"study_design_scores_gemma":[0.000018810579,0.000050679988,0.0013931451,0.00001850726,0.000027588992,0.00015979084,0.00007006009,0.9843566,0.0036494797,0.008153966,0.0020811458,0.000020264932],"about_ca_topic_score_codex":0.00735255,"about_ca_topic_score_gemma":0.009174095,"teacher_disagreement_score":0.00735255,"about_ca_system_score_codex":0.0009975054,"about_ca_system_score_gemma":0.0011188665,"threshold_uncertainty_score":0.014619529},"labels":[],"label_agreement":null},{"id":"W2128799871","doi":"10.1016/j.scico.2012.08.003","title":"Studying software evolution using topic models","year":2012,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":94,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Computer science; Structuring; Topic model; Software evolution; Source code; Software; Software development; Data science; Software maintenance; Task (project management); Software system; Code (set theory); Generative grammar; Software engineering; Information retrieval; Software construction; Artificial intelligence; Programming language","score_opus":0.05347133154518463,"score_gpt":0.3013226588468253,"score_spread":0.24785132730164067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128799871","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.655942,0.002896077,0.3319842,0.0028910728,0.00007529467,0.00009907073,0.00024229381,0.00032177652,0.0055480893],"genre_scores_gemma":[0.97195655,0.00094791426,0.024648491,0.00008693719,0.00011720856,0.000097115764,0.00026013728,0.00009288617,0.0017928516],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966814,0.0024480424,0.0001032021,0.0003427077,0.00024570356,0.00017888172],"domain_scores_gemma":[0.9218579,0.0719838,0.0020268979,0.001930685,0.0013629267,0.00083781866],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.006650657,0.0006665324,0.0010480109,0.0030670683,0.0009902626,0.0029002884,0.0014436508,0.001952693,0.0026030503],"category_scores_gemma":[0.052071422,0.0009069428,0.0015324674,0.0036039131,0.0009540666,0.008290567,0.0013152105,0.0019913653,0.00035262952],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007989627,0.0008040297,0.12463324,0.00083496305,0.0013271781,0.00061395223,0.006384851,0.34559876,0.006379063,0.3097241,0.006347086,0.19655386],"study_design_scores_gemma":[0.00006084682,0.000119183554,0.008352459,0.00003474368,0.0001738346,0.00016602178,0.00070504035,0.8728534,0.0008785827,0.114482336,0.0021438436,0.000029729357],"about_ca_topic_score_codex":0.0053309635,"about_ca_topic_score_gemma":0.0048311083,"teacher_disagreement_score":0.9933493,"about_ca_system_score_codex":0.001456497,"about_ca_system_score_gemma":0.0007379748,"threshold_uncertainty_score":0.035172462},"labels":[],"label_agreement":null},{"id":"W2129041774","doi":"","title":"Five Days of Empirical Software Engineering: the PASED Experience","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Victoria; Polytechnique Montréal","funders":"","keywords":"Empirical research; Supervisor; Software walkthrough; Plan (archaeology); Software Engineering Process Group; Social software engineering; Software engineering; Computer science; Personal software process; Software; Software peer review; Software development; Engineering management; Software construction; Engineering; Management","score_opus":0.022818383627743122,"score_gpt":0.2827374156048287,"score_spread":0.25991903197708555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129041774","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93483186,0.0009485897,0.020326879,0.015743582,0.0009108467,0.0007036733,0.0006439091,0.0008032393,0.025087463],"genre_scores_gemma":[0.9664806,0.0007002951,0.01669501,0.0020646022,0.00025360854,0.00030100162,0.0005107021,0.00021911577,0.012775056],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.9831266,0.009455807,0.0006716102,0.0016571188,0.0020697399,0.0030191059],"domain_scores_gemma":[0.93422437,0.020283615,0.0020270008,0.006265754,0.008578288,0.028621005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02460218,0.00097606855,0.0010615285,0.0014849274,0.0077902563,0.005939205,0.0035045892,0.002896493,0.00850872],"category_scores_gemma":[0.036429357,0.0011508606,0.0008399764,0.0013085179,0.007765329,0.0042391047,0.014720347,0.008408348,0.002131656],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010500811,0.009443901,0.039430466,0.00094167667,0.000106747466,0.010065806,0.6265124,0.0034907777,0.01485241,0.01346213,0.066024475,0.2146191],"study_design_scores_gemma":[0.00021045988,0.0037600272,0.03930613,0.0006792284,0.000039253573,0.0037415088,0.43827933,0.0050956924,0.0063751945,0.012222167,0.48996228,0.00032876205],"about_ca_topic_score_codex":0.0057895076,"about_ca_topic_score_gemma":0.020158425,"teacher_disagreement_score":0.02460218,"about_ca_system_score_codex":0.0063619637,"about_ca_system_score_gemma":0.008484554,"threshold_uncertainty_score":0.13011026},"labels":[],"label_agreement":null},{"id":"W2129065328","doi":"10.1145/2597073.2597100","title":"Finding patterns in static analysis alerts: improving actionable alert ranking","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":76,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Programmer; Computer science; Static analysis; Ranking (information retrieval); Software bug; Software; Data science; Software engineering; Data mining; Operating system; Information retrieval; Programming language","score_opus":0.014914032582510534,"score_gpt":0.26288812457693167,"score_spread":0.24797409199442114,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129065328","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34615627,0.0058577065,0.5919193,0.0029926999,0.00082468084,0.0009838792,0.009933943,0.033470422,0.007861056],"genre_scores_gemma":[0.68796545,0.0009536048,0.29038706,0.0005453078,0.0004090809,0.00024185568,0.014811177,0.0009861741,0.0037002116],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99236286,0.0014277633,0.0010263659,0.0013467383,0.0032479833,0.0005882459],"domain_scores_gemma":[0.96579885,0.014840038,0.0037568233,0.003847561,0.010316619,0.0014399996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044664396,0.002040202,0.0023607307,0.012349179,0.00088704185,0.0026930089,0.002499293,0.0019147823,0.0031367347],"category_scores_gemma":[0.035945166,0.0006023406,0.0009980967,0.0057771397,0.0005714586,0.0042846487,0.0017224278,0.001922986,0.0029725754],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012480282,0.0010119181,0.117176436,0.0007368881,0.00024474118,0.00030978528,0.00034696457,0.018097008,0.015177773,0.0028134112,0.035875373,0.80696166],"study_design_scores_gemma":[0.0002846148,0.0012578781,0.047258347,0.0002804418,0.00045989006,0.0012224535,0.0007819853,0.87734663,0.024769733,0.026567861,0.019543985,0.00022608005],"about_ca_topic_score_codex":0.006705629,"about_ca_topic_score_gemma":0.011645188,"teacher_disagreement_score":0.012349179,"about_ca_system_score_codex":0.00070250477,"about_ca_system_score_gemma":0.0024376113,"threshold_uncertainty_score":0.023621082},"labels":[],"label_agreement":null},{"id":"W2129182946","doi":"10.1007/s11219-007-9030-7","title":"Empirical studies to assess the understandability of data warehouse schemas using structural metrics","year":2007,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal; Université du Québec à Montréal","funders":"","keywords":"Computer science; Data mining; Metric (unit); Schema (genetic algorithms); Data warehouse; Similarity (geometry); Empirical research; Machine learning; Artificial intelligence; Statistics; Mathematics","score_opus":0.584142689272458,"score_gpt":0.5376677499733624,"score_spread":0.04647493929909563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129182946","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9971788,0.00011085559,0.0019188529,0.000066835135,0.0000031454,0.00006662227,0.000085018866,0.000012177549,0.00055779103],"genre_scores_gemma":[0.994705,0.00011240127,0.004469878,0.0000292937,0.000007247713,0.000075645,0.0004031491,0.000015745227,0.00018148597],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9783499,0.012099016,0.0032311694,0.0011236101,0.004687166,0.00050907873],"domain_scores_gemma":[0.35785428,0.57176244,0.031242324,0.012028494,0.025599087,0.0015134396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032877073,0.00055067195,0.0003442327,0.0039899335,0.00060045853,0.0015486703,0.0010766997,0.0010495356,0.0018656303],"category_scores_gemma":[0.26569083,0.00040672428,0.0007660653,0.004052996,0.0011400941,0.005919001,0.0010268776,0.0013581081,0.00024666215],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014029389,0.0061164177,0.90134746,0.0009957543,0.00057258626,0.00025676063,0.01747535,0.0036525235,0.0069306633,0.0027365454,0.00068218634,0.057830807],"study_design_scores_gemma":[0.00054807164,0.010917376,0.8933903,0.0005430387,0.0007444517,0.001103127,0.021806985,0.043310843,0.020101167,0.0035078556,0.0039140442,0.00011269851],"about_ca_topic_score_codex":0.0028477125,"about_ca_topic_score_gemma":0.0037708986,"teacher_disagreement_score":0.032877073,"about_ca_system_score_codex":0.0013482698,"about_ca_system_score_gemma":0.001245253,"threshold_uncertainty_score":0.17387265},"labels":[],"label_agreement":null},{"id":"W2129377409","doi":"10.1109/icsm.2015.7332456","title":"An empirical study of bugs in test code","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":108,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Software bug; Computer science; Code (set theory); Test (biology); Root cause; Regression testing; Code coverage; Programming language; Software; Reliability engineering; Software development; Engineering; Biology","score_opus":0.08033739344236439,"score_gpt":0.3817909511334575,"score_spread":0.3014535576910931,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129377409","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9941293,0.00048887223,0.0022172362,0.00028749762,0.000010106702,0.000070761715,0.0007957991,0.000034152505,0.0019664264],"genre_scores_gemma":[0.99817646,0.00014022451,0.00069117884,0.000048918733,0.000011204152,0.00006577915,0.00067153375,0.00001671928,0.00017797334],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9760105,0.010237865,0.002728961,0.0024894706,0.007650849,0.0008824648],"domain_scores_gemma":[0.48820207,0.36637172,0.10060881,0.015444699,0.025764415,0.0036083392],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012208534,0.00040329446,0.00032270467,0.005050029,0.00065101474,0.0015961857,0.001060989,0.00096518965,0.0020781157],"category_scores_gemma":[0.22117461,0.00050958426,0.00040232853,0.005326432,0.0021517498,0.0035646018,0.0014143727,0.0012566969,0.00042190193],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000121792466,0.00017158284,0.98198825,0.00019683666,0.00009717107,0.00017888799,0.0018942284,0.00049474847,0.00035450573,0.00071706984,0.0007480891,0.013036881],"study_design_scores_gemma":[0.000028603648,0.00046021264,0.9859254,0.00020728043,0.000048518734,0.0010832328,0.0036312507,0.003859547,0.0007349438,0.0008120622,0.0031847656,0.000024278657],"about_ca_topic_score_codex":0.0020833593,"about_ca_topic_score_gemma":0.0020335538,"teacher_disagreement_score":0.012208534,"about_ca_system_score_codex":0.00097270514,"about_ca_system_score_gemma":0.0007795447,"threshold_uncertainty_score":0.06456566},"labels":[],"label_agreement":null},{"id":"W2129412946","doi":"10.1109/wpc.2001.921719","title":"SHriMP views: an interactive environment for exploring Java programs","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Shrimp; Java; Computer science; Software; Visualization; Call graph; Software engineering; Programming language; Artificial intelligence; Ecology; Biology","score_opus":0.18028852038754978,"score_gpt":0.30142539667019597,"score_spread":0.12113687628264619,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129412946","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0150732035,0.00034352316,0.8526743,0.00027719856,0.00008048111,0.00021295346,0.0020613503,0.11322892,0.016048072],"genre_scores_gemma":[0.15320157,0.00091189996,0.8115822,0.00038346276,0.00007549814,0.0008351046,0.004498288,0.01523932,0.013272622],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996345,0.0001129136,0.00002338715,0.000056910045,0.0001318965,0.000040345836],"domain_scores_gemma":[0.9986106,0.0009831548,0.000047971782,0.00014497501,0.00007805139,0.0001351665],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00082239846,0.00080198806,0.00045827596,0.0011130329,0.0003827014,0.001368809,0.0013910255,0.00086103065,0.021122638],"category_scores_gemma":[0.0024465409,0.0005979144,0.0007429782,0.0005364755,0.00043568062,0.0028170713,0.0029494758,0.0015033023,0.003299503],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018368067,0.00044030332,0.0038447601,0.0014164317,0.00015246515,0.0018250258,0.004395796,0.015680399,0.15372369,0.0427993,0.16811743,0.6057675],"study_design_scores_gemma":[0.0006001807,0.0005715575,0.0053983824,0.0004968115,0.000114971255,0.0024075739,0.00073929597,0.19220988,0.08666104,0.04699545,0.6634403,0.00036452027],"about_ca_topic_score_codex":0.0009782913,"about_ca_topic_score_gemma":0.0021572022,"teacher_disagreement_score":0.021122638,"about_ca_system_score_codex":0.0002169727,"about_ca_system_score_gemma":0.00045868222,"threshold_uncertainty_score":0.07066226},"labels":[],"label_agreement":null},{"id":"W2129414299","doi":"10.1145/1370175.1370194","title":"Jigsaw","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Computer science; Code reuse; Jigsaw; Source code; KPI-driven code analysis; Context (archaeology); Code review; Code (set theory); Software engineering; Overhead (engineering); Software quality; Open source; Database; Software; Programming language; Software development; Engineering; Set (abstract data type)","score_opus":0.02905687703687048,"score_gpt":0.2535442770321326,"score_spread":0.22448739999526213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129414299","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007450643,0.0012786444,0.6335956,0.0007610563,0.00078809867,0.0005724918,0.0025822974,0.32583585,0.027135253],"genre_scores_gemma":[0.061510332,0.0017062834,0.79707515,0.0011759986,0.0002395821,0.0010427358,0.013794055,0.06200595,0.06144988],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977604,0.00029987938,0.00021223261,0.00046131233,0.0010661151,0.00020012388],"domain_scores_gemma":[0.9955537,0.0014332957,0.0003080085,0.0014196265,0.0009582032,0.00032715494],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026281804,0.0021004265,0.00078510353,0.0017841503,0.0013662981,0.0023481539,0.0028596516,0.0014906944,0.018792056],"category_scores_gemma":[0.007559802,0.0020268403,0.0017253902,0.001075261,0.00095073855,0.0031777523,0.0034553474,0.0030799198,0.010651204],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060989166,0.0003659735,0.0040335697,0.0022252102,0.0003007545,0.00065804506,0.001769199,0.009208268,0.042098332,0.033543956,0.35685927,0.54832757],"study_design_scores_gemma":[0.00026068787,0.00025854475,0.0024394235,0.00031523945,0.00017408005,0.00088430865,0.00012971245,0.030474478,0.032526974,0.03681818,0.8955442,0.00017426393],"about_ca_topic_score_codex":0.0026808353,"about_ca_topic_score_gemma":0.0045965877,"teacher_disagreement_score":0.018792056,"about_ca_system_score_codex":0.0005730166,"about_ca_system_score_gemma":0.0027398588,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2129594448","doi":"10.1109/icpc.2009.5090053","title":"A case for concept programs","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Executable; Comprehension; Realm; Program comprehension; Verifiable secret sharing; Meaning (existential); Programming language; Software engineering; Theoretical computer science; Artificial intelligence; Epistemology; Software","score_opus":0.0359440815137071,"score_gpt":0.30513473901879573,"score_spread":0.26919065750508864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129594448","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0355824,0.0028423725,0.71952057,0.061963607,0.00061141624,0.00022775274,0.00024001222,0.0008098095,0.17820206],"genre_scores_gemma":[0.76098275,0.0018153359,0.20936687,0.005857157,0.0008594787,0.00086570764,0.00022590064,0.00068868446,0.019338083],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99077517,0.0047913627,0.0002466301,0.0017041919,0.0017291077,0.00075349066],"domain_scores_gemma":[0.98016447,0.013947701,0.00068193953,0.0028000188,0.0016620996,0.0007438205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009896489,0.00068127335,0.00076315517,0.0018672165,0.0037919348,0.0072361133,0.002172178,0.006480451,0.0077017997],"category_scores_gemma":[0.026371362,0.0007121928,0.0014720157,0.0012629349,0.028728586,0.029319312,0.005470603,0.008394767,0.0010269915],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000051636002,0.0000031419397,0.00005102817,0.000010003201,0.0000013884181,0.00006661596,0.0005941984,0.00009232538,0.000059137572,0.9974164,0.0003698364,0.0013309654],"study_design_scores_gemma":[0.000013032175,0.000009952305,0.000057791804,0.000036480687,0.000004586925,0.00016321843,0.0003253717,0.0013018313,0.00025145305,0.96777725,0.030050704,0.000008335097],"about_ca_topic_score_codex":0.0025056812,"about_ca_topic_score_gemma":0.0009564465,"teacher_disagreement_score":0.009896489,"about_ca_system_score_codex":0.0026253336,"about_ca_system_score_gemma":0.0023560915,"threshold_uncertainty_score":0.052338243},"labels":[],"label_agreement":null},{"id":"W2129652343","doi":"10.1109/scam.2011.6","title":"What You See is What You Asked for: An Effort-Based Transformation of Code Analysis Tasks into Interactive Visualization Scenarios","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Computer science; Visualization; Task (project management); Context (archaeology); Process (computing); Code (set theory); Data visualization; Interactive visualization; Human–computer interaction; Task analysis; Transformation (genetics); Data mining; Programming language; Systems engineering","score_opus":0.03630839306163359,"score_gpt":0.32399045495334405,"score_spread":0.28768206189171047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129652343","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022609362,0.00004750546,0.9701287,0.0006843008,0.000016479786,0.0002790563,0.00015867777,0.0034632375,0.0026126422],"genre_scores_gemma":[0.21236883,0.00008295996,0.78461826,0.00014613944,0.000015228572,0.00037864278,0.0004072135,0.0005993974,0.0013833183],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9938275,0.003699882,0.00028081352,0.00067522016,0.001255528,0.0002610104],"domain_scores_gemma":[0.9877995,0.0075502032,0.0008853618,0.0017911854,0.0015518939,0.00042177888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045071323,0.0014833206,0.0005822351,0.0015844685,0.00066977635,0.0032824262,0.001664776,0.0016949945,0.0041974834],"category_scores_gemma":[0.019147173,0.00083325774,0.0011952963,0.00080960355,0.0012523159,0.003087105,0.0023807802,0.0016121768,0.001563948],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017705283,0.0012815187,0.011069576,0.0011279042,0.00023963906,0.0014697483,0.013720298,0.093162194,0.11356025,0.08302238,0.012006198,0.6675698],"study_design_scores_gemma":[0.00020052903,0.000651096,0.006131423,0.00032457212,0.00013078148,0.001008446,0.0035622127,0.80707043,0.06514299,0.08200367,0.033471514,0.00030232323],"about_ca_topic_score_codex":0.001446664,"about_ca_topic_score_gemma":0.0019404894,"teacher_disagreement_score":0.0045071323,"about_ca_system_score_codex":0.00060053583,"about_ca_system_score_gemma":0.0010972432,"threshold_uncertainty_score":0.023836315},"labels":[],"label_agreement":null},{"id":"W2129879174","doi":"10.1109/icpc.2006.45","title":"Summarizing the Content of Large Traces to Facilitate the Understanding of the Behaviour of a Software System","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":139,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Automatic summarization; TRACE (psycholinguistics); Computer science; Sequence diagram; Unified Modeling Language; Metric (unit); Rank (graph theory); Software; Software system; Representation (politics); Key (lock); Software engineering; Information retrieval; Data mining; Programming language; Operating system","score_opus":0.09778230045110695,"score_gpt":0.25917293022990806,"score_spread":0.16139062977880111,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129879174","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05197915,0.0011086692,0.92923105,0.00045676008,0.00013918884,0.0006908939,0.005047963,0.009392761,0.0019535574],"genre_scores_gemma":[0.15067473,0.0010537205,0.8310403,0.00009836448,0.00015603744,0.00051286793,0.013110524,0.001472207,0.0018812421],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99789256,0.0005382694,0.00029292676,0.00035354795,0.0008316995,0.0000910543],"domain_scores_gemma":[0.9764227,0.014173716,0.0023055566,0.0025872958,0.004173543,0.00033717445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002534118,0.0018579334,0.0009575511,0.009685599,0.00083712826,0.002155255,0.0010353768,0.0009034442,0.0029955977],"category_scores_gemma":[0.020235647,0.0004733742,0.000816919,0.004638078,0.0004885305,0.003256465,0.0010716772,0.0011709672,0.0013118084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008558724,0.00030975215,0.01235413,0.0029031762,0.00027102712,0.0013065913,0.0046083955,0.030161401,0.082690686,0.010784521,0.012335616,0.8414187],"study_design_scores_gemma":[0.00025352443,0.0016763193,0.051033236,0.0020561768,0.0013035053,0.0034977533,0.005886679,0.46491018,0.1820237,0.12031867,0.16633311,0.00070701115],"about_ca_topic_score_codex":0.0028343513,"about_ca_topic_score_gemma":0.0031318006,"teacher_disagreement_score":0.009685599,"about_ca_system_score_codex":0.0006162754,"about_ca_system_score_gemma":0.0013309808,"threshold_uncertainty_score":0.013401866},"labels":[],"label_agreement":null},{"id":"W2129903536","doi":"10.1109/ijcnn.2006.246954","title":"Fuzzy Clustering of Open-Source Software Quality Data: A Case Study of Mozilla","year":2006,"lang":"en","type":"article","venue":"The 2006 IEEE International Joint Conference on Neural Network Proceedings","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Data mining; Software quality; Fuzzy logic; Software; Cluster analysis; Software metric; Ranking (information retrieval); Quality (philosophy); Software sizing; Metric (unit); Artificial intelligence; Software construction; Software system; Software development; Operating system; Engineering","score_opus":0.13971397060377239,"score_gpt":0.3598691079223464,"score_spread":0.220155137318574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129903536","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98431814,0.0004818055,0.011435561,0.0006128681,0.000015199126,0.000120428514,0.0013307937,0.00039775143,0.0012874199],"genre_scores_gemma":[0.95619446,0.00023940687,0.03710746,0.000077256984,0.000017343848,0.00012481285,0.004648414,0.00015576006,0.0014350015],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99701893,0.00075850164,0.00020394148,0.00057062256,0.0012192853,0.00022865403],"domain_scores_gemma":[0.9876923,0.006318819,0.0015035226,0.0014248051,0.0024988751,0.0005616987],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003705156,0.00046582212,0.00072507275,0.0059329923,0.0016794133,0.0015456363,0.0013401966,0.0015262876,0.0003570806],"category_scores_gemma":[0.01524006,0.00027031236,0.00095413334,0.0069112345,0.00096120324,0.0010501822,0.0011254175,0.0008049131,0.00019876384],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002209289,0.0015417543,0.5000862,0.0014074087,0.0010233903,0.011210754,0.015978143,0.10784637,0.03561741,0.00854288,0.02090701,0.29362938],"study_design_scores_gemma":[0.00025288615,0.00046346235,0.6373758,0.00024276144,0.00025039847,0.0034260668,0.006231265,0.29446346,0.02592441,0.0055094683,0.025618212,0.00024184727],"about_ca_topic_score_codex":0.053166192,"about_ca_topic_score_gemma":0.06267132,"teacher_disagreement_score":0.053166192,"about_ca_system_score_codex":0.0028136268,"about_ca_system_score_gemma":0.0010470144,"threshold_uncertainty_score":0.10571343},"labels":[],"label_agreement":null},{"id":"W2130011984","doi":"10.1109/qsic.2009.69","title":"Quality of the Source Code for Design and Architecture Recovery Techniques: Utilities are the Problem","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Architecture; Source code; Software architecture; Legacy system; Software engineering; Reverse engineering; Code (set theory); Software maintenance; Quality (philosophy); Software quality; Software system; Software; Reliability engineering; Software development; Engineering; Programming language","score_opus":0.041638038078125376,"score_gpt":0.29297993784555676,"score_spread":0.2513418997674314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130011984","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08769503,0.0024247204,0.8810459,0.008980292,0.0003648303,0.00044327648,0.00049042795,0.013272219,0.0052833282],"genre_scores_gemma":[0.52113104,0.0015105987,0.45686263,0.0023521918,0.0002250825,0.0004507955,0.0014824213,0.009001158,0.0069841747],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95647347,0.011885252,0.003238713,0.0034497757,0.024242582,0.00071020925],"domain_scores_gemma":[0.5684833,0.20023876,0.037294414,0.12489821,0.06716677,0.0019185806],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028955093,0.0009934002,0.0012375944,0.0032589517,0.0015321544,0.0042986716,0.0032348006,0.0028434477,0.002234184],"category_scores_gemma":[0.25981835,0.0015146385,0.0008513522,0.0030325684,0.0032508965,0.0070656086,0.002390092,0.0040009525,0.0022331998],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009323756,0.0005460584,0.035634268,0.0018490569,0.00020275988,0.0011515614,0.004692927,0.022661138,0.0344797,0.05509379,0.018176313,0.82457995],"study_design_scores_gemma":[0.0005560347,0.0011753497,0.046410315,0.0035646919,0.00057459594,0.006967898,0.0024333068,0.3903787,0.21430513,0.17240031,0.16070792,0.0005257544],"about_ca_topic_score_codex":0.0032230695,"about_ca_topic_score_gemma":0.0021081616,"teacher_disagreement_score":0.9710449,"about_ca_system_score_codex":0.0020796906,"about_ca_system_score_gemma":0.0047074966,"threshold_uncertainty_score":0.15313095},"labels":[],"label_agreement":null},{"id":"W2130273245","doi":"10.5381/jot.2008.7.6.a1","title":"Revisiting Class Cohesion: An empirical investigation on several systems.","year":2008,"lang":"en","type":"article","venue":"The Journal of Object Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Cohesion (chemistry); Class (philosophy); Computer science; Empirical research; Artificial intelligence; Mathematics; Chemistry; Statistics","score_opus":0.04793237885926566,"score_gpt":0.29829285571453223,"score_spread":0.2503604768552666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130273245","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.989593,0.00045427767,0.007850839,0.00012961627,0.00001607082,0.000103743056,0.00026571625,0.000079813304,0.0015068694],"genre_scores_gemma":[0.99581486,0.00010542061,0.0032383048,0.000022489052,0.000013568254,0.00009482684,0.00042403606,0.0000286706,0.00025783567],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9881754,0.0046758954,0.001154932,0.0018569123,0.0037881297,0.00034880463],"domain_scores_gemma":[0.7639451,0.18447742,0.020513183,0.012380743,0.016560603,0.0021229591],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0110585885,0.0004321931,0.00043582017,0.0047779833,0.0010386518,0.0011301743,0.0009368289,0.0007314438,0.0013816555],"category_scores_gemma":[0.11794995,0.0002618793,0.0005809463,0.003931923,0.0017060669,0.0029474476,0.0016256155,0.0010262389,0.00032033582],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003845786,0.000710424,0.86045635,0.00084959454,0.0005329689,0.00058467843,0.013091207,0.005251573,0.0058701807,0.002431052,0.0022711908,0.10756616],"study_design_scores_gemma":[0.00006180026,0.0017487045,0.95412827,0.000141571,0.00021485865,0.0012174534,0.0061729876,0.022223074,0.0049067535,0.0023323589,0.0067755156,0.00007669281],"about_ca_topic_score_codex":0.001680913,"about_ca_topic_score_gemma":0.0023108886,"teacher_disagreement_score":0.0110585885,"about_ca_system_score_codex":0.0009921055,"about_ca_system_score_gemma":0.00055639294,"threshold_uncertainty_score":0.058484077},"labels":[],"label_agreement":null},{"id":"W2130373880","doi":"10.1109/coginf.2002.1039314","title":"Modeling comprehension processes in software development","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Program comprehension; Software development; Software analytics; Abstraction; Software engineering; Software; Comprehension; Software development process; A priori and a posteriori; Software maintenance; Cognition; Software system; Human–computer interaction; Programming language","score_opus":0.03591835732360858,"score_gpt":0.2654689460659842,"score_spread":0.2295505887423756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130373880","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18305974,0.0018856998,0.7846919,0.0017164132,0.00007447231,0.00032886342,0.000209067,0.00051471917,0.027519181],"genre_scores_gemma":[0.8369489,0.0014794791,0.15661386,0.00011825762,0.00007376839,0.000605309,0.00023657132,0.00010241296,0.0038215427],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974233,0.0016527151,0.0001136195,0.00032383087,0.00028628955,0.00020014976],"domain_scores_gemma":[0.9838504,0.013418788,0.0010040776,0.00065570656,0.0007800176,0.00029093557],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038706022,0.0011461311,0.00044954382,0.0020415657,0.00091556844,0.003815844,0.0014188469,0.0025202096,0.0029965008],"category_scores_gemma":[0.025285814,0.0006896523,0.00094189955,0.0018282423,0.0024160543,0.0067202896,0.001967902,0.0016601479,0.00065859384],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033096483,0.00051235664,0.0136190085,0.0004185438,0.00013742551,0.0005615306,0.01498077,0.4696763,0.0032035888,0.4267005,0.0013875944,0.06847143],"study_design_scores_gemma":[0.000072674244,0.0001125824,0.0021604735,0.000076805074,0.000055892437,0.000083654806,0.00076827744,0.7084103,0.0006422591,0.2828856,0.004691357,0.00004010155],"about_ca_topic_score_codex":0.010620877,"about_ca_topic_score_gemma":0.005889642,"teacher_disagreement_score":0.010620877,"about_ca_system_score_codex":0.0020479574,"about_ca_system_score_gemma":0.002223756,"threshold_uncertainty_score":0.021118164},"labels":[],"label_agreement":null},{"id":"W2130419774","doi":"10.1109/icpc.2008.35","title":"An Approach for Mapping Features to Code Based on Static and Dynamic Analysis","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Feature (linguistics); TRACE (psycholinguistics); Feature model; Tracing; Static analysis; Data mining; Component (thermodynamics); Domain (mathematical analysis); Source code; Relevance (law); Dependency (UML); Feature extraction; Artificial intelligence; Software; Programming language","score_opus":0.024194017413483267,"score_gpt":0.2872019440593998,"score_spread":0.26300792664591655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130419774","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021698384,0.0000604046,0.9919321,0.000131202,0.000026745573,0.00010758317,0.000073847106,0.0044409963,0.0010572114],"genre_scores_gemma":[0.044629857,0.00014078623,0.9512268,0.00011631531,0.000031007967,0.00018182746,0.00029169748,0.0009642671,0.0024176163],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9969215,0.0005082987,0.00018385884,0.00074910803,0.0014880924,0.00014905652],"domain_scores_gemma":[0.9930728,0.0019188006,0.00076203415,0.0025953685,0.0015275023,0.00012350433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023073119,0.0015913775,0.0009970673,0.00590621,0.0012019748,0.0026452031,0.0025129726,0.0014000601,0.0030032236],"category_scores_gemma":[0.009815474,0.0012567814,0.0017336608,0.0029641576,0.0016519238,0.0037427242,0.0023220852,0.002641095,0.001572674],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013820502,0.0003611221,0.0040520662,0.00051171944,0.00019352179,0.00040769513,0.0013173192,0.027953112,0.046253406,0.05351845,0.0072109792,0.85808235],"study_design_scores_gemma":[0.000088878565,0.00046627125,0.0048526055,0.00028965136,0.00043198242,0.0018327003,0.0006121309,0.6337206,0.10922779,0.13968895,0.10841788,0.00037063903],"about_ca_topic_score_codex":0.0054834625,"about_ca_topic_score_gemma":0.007990058,"teacher_disagreement_score":0.00590621,"about_ca_system_score_codex":0.0014315465,"about_ca_system_score_gemma":0.0031005575,"threshold_uncertainty_score":0.012202382},"labels":[],"label_agreement":null},{"id":"W2130602377","doi":"10.1109/mtd.2015.7332619","title":"Detecting and quantifying different types of self-admitted technical Debt","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":163,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Technical debt; Debt; Computer science; Documentation; Debt levels and flows; Risk analysis (engineering); Internal debt; Business; Finance; Software; Software development","score_opus":0.05419014256578286,"score_gpt":0.30355784438373046,"score_spread":0.2493677018179476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130602377","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9651109,0.0011952532,0.02320269,0.00041199796,0.00009238478,0.00022912296,0.004475361,0.0012387782,0.0040434357],"genre_scores_gemma":[0.9366555,0.00075433357,0.045532633,0.00026048793,0.00011119334,0.0002643761,0.011031189,0.00046379276,0.0049264263],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9920844,0.0016484306,0.0011815863,0.0011034813,0.003497164,0.00048500736],"domain_scores_gemma":[0.8439914,0.07917387,0.033717632,0.00818422,0.032411136,0.002521701],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004921887,0.0008206325,0.0005120829,0.008852834,0.0007882685,0.0016255783,0.00084265234,0.0012111963,0.000968486],"category_scores_gemma":[0.07128883,0.00037812497,0.00040661497,0.0048787314,0.0005542617,0.0030355772,0.0015076826,0.0009283747,0.00075526274],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005423414,0.00028971804,0.7798773,0.0018499317,0.0001706828,0.0014441253,0.006907241,0.0022542586,0.024720723,0.0018275853,0.010303096,0.16981319],"study_design_scores_gemma":[0.000057536563,0.0005471369,0.83101463,0.0012482104,0.00024516214,0.004328971,0.010620448,0.077989124,0.029002298,0.0049494724,0.039758183,0.0002388375],"about_ca_topic_score_codex":0.0037853676,"about_ca_topic_score_gemma":0.007228362,"teacher_disagreement_score":0.008852834,"about_ca_system_score_codex":0.0008845702,"about_ca_system_score_gemma":0.0011844503,"threshold_uncertainty_score":0.026029706},"labels":[],"label_agreement":null},{"id":"W2131077449","doi":"10.1109/rev.2006.10","title":"Visualizing non-functional requirements","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Functional requirement; Visualization; Non-functional requirement; Scheme (mathematics); Architecture; Quality (philosophy); Unified Modeling Language; Software engineering; Programming language; Human–computer interaction; Artificial intelligence; Software development; Software","score_opus":0.02976248698628348,"score_gpt":0.2921429460850965,"score_spread":0.262380459098813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131077449","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039094187,0.00055784837,0.93391263,0.0011939164,0.00010525656,0.00014747196,0.0011796582,0.0044618617,0.01934714],"genre_scores_gemma":[0.41383842,0.0009811281,0.57539576,0.00015907285,0.000060564787,0.00022481346,0.0018784066,0.0011812489,0.0062805284],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986519,0.00059250597,0.00009434616,0.000106163076,0.0004717308,0.00008326972],"domain_scores_gemma":[0.9954,0.0024144545,0.00034258075,0.0006424955,0.0009849432,0.00021562587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022012438,0.0013404423,0.00040065157,0.0026470656,0.00041548448,0.0027350413,0.0008767362,0.0010759307,0.0076500326],"category_scores_gemma":[0.009074522,0.00035936746,0.00065570005,0.0012705475,0.0004992264,0.0029685814,0.0016886547,0.000965743,0.0010997007],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046051265,0.0002161702,0.005785553,0.0021842956,0.00015294069,0.0012540029,0.013073929,0.07686727,0.10687768,0.36637047,0.030731054,0.39602613],"study_design_scores_gemma":[0.00014774059,0.0004387781,0.0123504205,0.0006419395,0.00018350815,0.002011804,0.004526537,0.42301762,0.08036329,0.25106925,0.2249772,0.0002718945],"about_ca_topic_score_codex":0.001803163,"about_ca_topic_score_gemma":0.001406485,"teacher_disagreement_score":0.0076500326,"about_ca_system_score_codex":0.0005427806,"about_ca_system_score_gemma":0.0006399793,"threshold_uncertainty_score":0.02559197},"labels":[],"label_agreement":null},{"id":"W2131112324","doi":"10.1109/wcre.2009.51","title":"An Empirical Study on Inconsistent Changes to Code Clones at Release Level","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software evolution; Computer science; Software; Code (set theory); Perspective (graphical); Software quality; Empirical research; Quality (philosophy); Software development; Software release life cycle; Software engineering; Open source software; Code review; Software construction; Programming language; Artificial intelligence; Statistics; Mathematics","score_opus":0.11304633405163557,"score_gpt":0.3774895446032434,"score_spread":0.26444321055160785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131112324","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9980399,0.00016656258,0.0005819966,0.000172761,0.000004842846,0.000025067311,0.00003772082,0.000005613633,0.0009655273],"genre_scores_gemma":[0.99909174,0.00009617337,0.0004210434,0.00010053251,0.000010271452,0.000028686321,0.00007276945,0.000008946916,0.00016988184],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9740054,0.014044038,0.0021796322,0.002091968,0.0069570923,0.0007218346],"domain_scores_gemma":[0.2519299,0.5780786,0.12470529,0.019108467,0.02293366,0.0032440028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019535009,0.0002963089,0.00033730266,0.0020569428,0.00086730224,0.0018270531,0.0013121314,0.0013648174,0.0020683086],"category_scores_gemma":[0.3148227,0.00040717292,0.0002904919,0.0021185947,0.0023009758,0.00335118,0.001308139,0.0025244192,0.0003098545],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041383752,0.0008082349,0.97091246,0.00015828462,0.00008245135,0.0004510378,0.00747525,0.0003489447,0.00074857794,0.00080176536,0.00036663833,0.017432328],"study_design_scores_gemma":[0.000039101404,0.0011412936,0.98583525,0.00011908146,0.00006799955,0.0009166403,0.007542721,0.0015313047,0.00068735634,0.0005401462,0.0015527108,0.000026452262],"about_ca_topic_score_codex":0.0013537143,"about_ca_topic_score_gemma":0.0014304331,"teacher_disagreement_score":0.019535009,"about_ca_system_score_codex":0.00080530817,"about_ca_system_score_gemma":0.00074635324,"threshold_uncertainty_score":0.103312194},"labels":[],"label_agreement":null},{"id":"W2131153512","doi":"10.1109/icsm.2005.31","title":"Comparison of clustering algorithms in the context of software evolution","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":128,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cluster analysis; Computer science; Data mining; Software; Partition (number theory); Software system; Context (archaeology); Stability (learning theory); Cluster (spacecraft); Algorithm; CURE data clustering algorithm; Software maintenance; Fuzzy clustering; Machine learning; Mathematics; Programming language","score_opus":0.035006692003917046,"score_gpt":0.3216598830322398,"score_spread":0.28665319102832276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131153512","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45986578,0.0052600144,0.51954335,0.0011375678,0.00037453804,0.0005199583,0.0005811862,0.0025706643,0.010146905],"genre_scores_gemma":[0.55997086,0.0016647977,0.43521255,0.00010785014,0.00009339216,0.00025365665,0.001383643,0.00031067754,0.0010026623],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.989593,0.0051275794,0.00085206924,0.0010859144,0.0029782073,0.00036314444],"domain_scores_gemma":[0.9617875,0.024586389,0.0014380601,0.0029229592,0.008870021,0.00039506602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010598188,0.0009884966,0.0011170603,0.008656469,0.0017152302,0.00253873,0.0020545765,0.0019473648,0.0008226697],"category_scores_gemma":[0.049207788,0.0003910107,0.001046826,0.007878119,0.00087746535,0.0029865552,0.0012463648,0.00093095755,0.0003918191],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015618547,0.00041539877,0.027922979,0.0008416756,0.0010233794,0.00013834423,0.0016882941,0.29817557,0.004976518,0.022651536,0.0052206335,0.63538384],"study_design_scores_gemma":[0.00015019803,0.0003785161,0.018801695,0.00010098177,0.00019913875,0.00031047466,0.001213582,0.94608295,0.00775472,0.019596787,0.0052846093,0.00012641936],"about_ca_topic_score_codex":0.007712364,"about_ca_topic_score_gemma":0.007527138,"teacher_disagreement_score":0.010598188,"about_ca_system_score_codex":0.0023546955,"about_ca_system_score_gemma":0.0020570925,"threshold_uncertainty_score":0.056049228},"labels":[],"label_agreement":null},{"id":"W2131491818","doi":"10.1109/icpc.2009.5090048","title":"Vector space analysis of software clones","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Representation (politics); Code (set theory); Source code; Software; Similarity (geometry); Vector space; clone (Java method); Matrix representation; Space (punctuation); Algorithm; Artificial intelligence; Mathematics; Programming language; Pure mathematics; Group (periodic table); Physics","score_opus":0.013179686168611579,"score_gpt":0.26747614420694815,"score_spread":0.25429645803833656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131491818","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07649884,0.00033999424,0.92022717,0.000113236594,0.000030470996,0.000074235606,0.00015729369,0.00094479864,0.0016138877],"genre_scores_gemma":[0.66945666,0.000516133,0.32461712,0.00007384367,0.00006620207,0.0002370127,0.000925804,0.00027780875,0.0038293723],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984517,0.000289223,0.00007862681,0.00025715813,0.0007829081,0.00014027803],"domain_scores_gemma":[0.99488527,0.0018526455,0.0006350154,0.00045609643,0.0020174496,0.00015356283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010751172,0.000650455,0.00061922043,0.005266792,0.000681703,0.0015332388,0.0008390629,0.0006195303,0.0021822553],"category_scores_gemma":[0.009866943,0.00028122318,0.0007365254,0.004013412,0.00094289175,0.0022168146,0.000963696,0.00077686505,0.0005171884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047252802,0.00013855238,0.017494021,0.0002795202,0.00015134638,0.00042080611,0.0014163995,0.14477246,0.025196124,0.1320843,0.0027573556,0.6748166],"study_design_scores_gemma":[0.00003141127,0.00022865279,0.0093028955,0.000047934947,0.00006248307,0.00029749764,0.0004267761,0.8819296,0.009552187,0.090691894,0.0073548043,0.00007392685],"about_ca_topic_score_codex":0.0072641117,"about_ca_topic_score_gemma":0.002807729,"teacher_disagreement_score":0.0072641117,"about_ca_system_score_codex":0.0009359662,"about_ca_system_score_gemma":0.0010657223,"threshold_uncertainty_score":0.014443636},"labels":[],"label_agreement":null},{"id":"W2131584723","doi":"10.1109/wcre.2006.7","title":"A Service-Oriented Componentization Framework for Java Software Systems","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Java; Software engineering; Programming language; Operating system","score_opus":0.015453489901974206,"score_gpt":0.25724330812047996,"score_spread":0.24178981821850576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131584723","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00059879967,0.00014452307,0.9938538,0.000111987974,0.000032975906,0.00015328985,0.000046967903,0.003347855,0.0017099252],"genre_scores_gemma":[0.020346569,0.00046432097,0.9744475,0.000109719396,0.000041183484,0.00047046362,0.0003585699,0.00096840196,0.0027932671],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9986903,0.0002919793,0.00015925421,0.00015383045,0.00056515634,0.00013944074],"domain_scores_gemma":[0.99939096,0.0001789542,0.00005976856,0.000112495545,0.00018235309,0.000075314834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022265448,0.0012746669,0.0007678695,0.0015058038,0.0011470021,0.0024615957,0.002464303,0.0013068987,0.0029312186],"category_scores_gemma":[0.002230527,0.0010733232,0.0019710737,0.0014023028,0.0012808138,0.0020178873,0.0014592494,0.002779439,0.002046916],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011400385,0.00020217699,0.00066394417,0.00068507483,0.000099797406,0.0008282223,0.0012396824,0.060598616,0.021256093,0.6749503,0.01944313,0.2199189],"study_design_scores_gemma":[0.00014136366,0.00014345965,0.0004909678,0.00030088995,0.00013005643,0.0010733663,0.00014740136,0.38907248,0.01293614,0.16080438,0.43460256,0.00015688578],"about_ca_topic_score_codex":0.010234073,"about_ca_topic_score_gemma":0.010923234,"teacher_disagreement_score":0.010234073,"about_ca_system_score_codex":0.0012777653,"about_ca_system_score_gemma":0.0031448775,"threshold_uncertainty_score":0.020348966},"labels":[],"label_agreement":null},{"id":"W2131778596","doi":"10.1109/scam.2009.27","title":"A Metric Extraction Framework Based on a High-Level Description Language","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal","funders":"","keywords":"Computer science; Metric (unit); Programming language; Focus (optics); Code (set theory); Implementation; Software metric; Theoretical computer science; Source code; Object-oriented programming; Black box; Algorithm; Software; Software development; Artificial intelligence; Software quality","score_opus":0.031236502867829294,"score_gpt":0.29958171404208367,"score_spread":0.26834521117425436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131778596","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00022253505,0.00008439738,0.9966737,0.000102054444,0.000016944767,0.00012920627,0.0002564758,0.0019766618,0.00053789583],"genre_scores_gemma":[0.0073259175,0.0002643274,0.9891806,0.00009581999,0.000025308469,0.00038257844,0.0013467981,0.0006006091,0.0007779932],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913304,0.0019610592,0.0018294296,0.0010467309,0.00352744,0.0003048961],"domain_scores_gemma":[0.9923413,0.0024778026,0.0009362149,0.0015887878,0.0023233346,0.0003326649],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009221191,0.0021504068,0.0017339902,0.007215917,0.0013427864,0.0062889275,0.0041849376,0.0018190182,0.0034122162],"category_scores_gemma":[0.011298296,0.0016752246,0.0039512315,0.0050070356,0.0019217316,0.0069196136,0.0031578601,0.0042771585,0.0025708168],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000856924,0.00018471011,0.0013036479,0.0018270857,0.0002119017,0.0008887775,0.0010022122,0.036462557,0.01357204,0.7072009,0.018576583,0.218684],"study_design_scores_gemma":[0.00008589787,0.0002564065,0.00094007875,0.0009086075,0.00023461806,0.0017177175,0.0002922298,0.26889595,0.021835102,0.2675807,0.43689683,0.0003558545],"about_ca_topic_score_codex":0.007099473,"about_ca_topic_score_gemma":0.0068298196,"teacher_disagreement_score":0.009221191,"about_ca_system_score_codex":0.0026974652,"about_ca_system_score_gemma":0.006660298,"threshold_uncertainty_score":0.04876691},"labels":[],"label_agreement":null},{"id":"W2131861583","doi":"10.1109/icsm.2008.4658104","title":"Task articulation in software maintenance: Integrating source code annotations with an issue tracking system","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Task management; Software engineering; Task (project management); Software maintenance; Software development; Software construction; Source code; Software system; Variety (cybernetics); Backporting; Software; Human–computer interaction; Systems engineering; Programming language; Artificial intelligence; Engineering","score_opus":0.01934763698472294,"score_gpt":0.2568629803440977,"score_spread":0.23751534335937477,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131861583","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015054239,0.00018800845,0.94313556,0.0007146266,0.0002033562,0.00025434533,0.00009691562,0.038364958,0.001987982],"genre_scores_gemma":[0.092806295,0.00016368522,0.90195554,0.00021763354,0.00014559674,0.00020259015,0.00037768146,0.0019489606,0.00218199],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9868623,0.005217579,0.001707449,0.0017971495,0.004018301,0.00039713766],"domain_scores_gemma":[0.9153509,0.0454783,0.009910449,0.01847546,0.008318192,0.00246668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022139877,0.0011740542,0.0010981718,0.004752204,0.0020769255,0.00574218,0.0040621744,0.003141377,0.002498971],"category_scores_gemma":[0.08631521,0.002050565,0.00101835,0.0034567302,0.002426517,0.010038513,0.005789755,0.005048625,0.0018929077],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010264363,0.001646822,0.013906475,0.00075156294,0.00017494279,0.0010414931,0.009127764,0.017792098,0.032883722,0.029130049,0.018103497,0.87441516],"study_design_scores_gemma":[0.00062348926,0.00093889114,0.010801069,0.0010825241,0.0007533849,0.001923267,0.0015973293,0.6522159,0.10518444,0.057966053,0.16596341,0.00095028646],"about_ca_topic_score_codex":0.0035174065,"about_ca_topic_score_gemma":0.0037305893,"teacher_disagreement_score":0.022139877,"about_ca_system_score_codex":0.0011495289,"about_ca_system_score_gemma":0.004060434,"threshold_uncertainty_score":0.11708826},"labels":[],"label_agreement":null},{"id":"W2132194716","doi":"10.1109/wpc.2001.921716","title":"Navigation and comprehension of programs by novice programmers","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Programmer; Program comprehension; Computer science; Comprehension; Control flow; Human–computer interaction; Information flow; Process (computing); Mental representation; Representation (politics); Control (management); Programming language; Software engineering; Multimedia; Artificial intelligence; Software; Cognition; Software system","score_opus":0.02216530389166788,"score_gpt":0.24837855725783592,"score_spread":0.22621325336616804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132194716","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9957228,0.00015235113,0.0022859506,0.000057351455,0.0000026370626,0.000016923204,0.000018806173,0.0000742229,0.001668904],"genre_scores_gemma":[0.99580306,0.00020496073,0.0024756521,0.000058875037,0.000005622127,0.000017033513,0.000116174924,0.000029348534,0.0012894557],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9987513,0.0005787694,0.00007090253,0.00019487651,0.00030945163,0.0000946916],"domain_scores_gemma":[0.97352284,0.018650226,0.0029009522,0.0014679902,0.0025180338,0.0009399802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014293676,0.00041331514,0.00027814152,0.0006051997,0.0002486228,0.0016305652,0.0003820103,0.0005823619,0.002061767],"category_scores_gemma":[0.03728349,0.0002368165,0.00029747337,0.00020434688,0.00053069263,0.0015762713,0.00074498146,0.0008895543,0.00038826742],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022853361,0.002593665,0.40661013,0.0010740014,0.0002926144,0.0012579696,0.07461908,0.0060108863,0.10742692,0.0020834245,0.0045669535,0.39117897],"study_design_scores_gemma":[0.00024162843,0.009538308,0.86963445,0.00040821257,0.00027507052,0.002352445,0.025202045,0.027916774,0.04597209,0.0063269534,0.011886991,0.00024504997],"about_ca_topic_score_codex":0.0014033248,"about_ca_topic_score_gemma":0.001399014,"teacher_disagreement_score":0.002061767,"about_ca_system_score_codex":0.00022680147,"about_ca_system_score_gemma":0.00033805467,"threshold_uncertainty_score":0.0075592995},"labels":[],"label_agreement":null},{"id":"W2132333108","doi":"10.1109/icsm.2005.6","title":"A category-theoretic approach to syntactic software merging","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Merge (version control); Theoretical computer science; Software; Exploit; Morphism; Software system; Data mining; Artificial intelligence; Programming language; Information retrieval; Mathematics","score_opus":0.015656138586761882,"score_gpt":0.25363952770926274,"score_spread":0.23798338912250086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132333108","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0043086684,0.0003844682,0.97701025,0.00079663214,0.000073966236,0.000055403278,0.00005838825,0.00019876526,0.017113555],"genre_scores_gemma":[0.23928359,0.0006945093,0.75073713,0.00053006824,0.0002954103,0.000298649,0.00021048395,0.00015228159,0.007797895],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99658406,0.0010475457,0.00024699347,0.0005370822,0.0013339997,0.00025028392],"domain_scores_gemma":[0.9969007,0.0011750739,0.0002584905,0.00067229435,0.0008295609,0.00016398737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003408889,0.0006478828,0.00062721554,0.005817659,0.00322619,0.004046217,0.0029792208,0.0019703722,0.0042132726],"category_scores_gemma":[0.0051199724,0.0006863761,0.0026919944,0.0040134713,0.0076187225,0.009028061,0.004969226,0.003392352,0.0006479019],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000037020782,0.000007525451,0.00008126271,0.000027038826,0.000009519005,0.000043110293,0.00022365934,0.0017127465,0.00030714323,0.99120647,0.00032875323,0.006049025],"study_design_scores_gemma":[0.000006531906,0.000017056931,0.00009600121,0.000019494124,0.000014993147,0.00015781021,0.00012766411,0.015299063,0.00066165585,0.971464,0.012116749,0.000019038584],"about_ca_topic_score_codex":0.0034799105,"about_ca_topic_score_gemma":0.003401416,"teacher_disagreement_score":0.005817659,"about_ca_system_score_codex":0.0035430929,"about_ca_system_score_gemma":0.0018493949,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2132387930","doi":"10.1080/03081070903029253","title":"Synergistic verification and validation of systems and software engineering models","year":2009,"lang":"en","type":"article","venue":"International Journal of General Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; Concordia University","funders":"","keywords":"Computer science; Systems Modeling Language; Unified Modeling Language; Software engineering; Verification and validation; Software verification; Sequence diagram; Formal verification; Software; Software system; Programming language; Data mining; Systems engineering; Software construction; Engineering","score_opus":0.019569621150831835,"score_gpt":0.2599739429218604,"score_spread":0.24040432177102855,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132387930","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050768936,0.00019576894,0.99214387,0.00015500338,0.000030972013,0.00011199132,0.000033601627,0.00064160017,0.0016103438],"genre_scores_gemma":[0.19636402,0.0005896566,0.7999482,0.00016178162,0.0000494975,0.00066532363,0.00034694895,0.0002934262,0.0015811602],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.951293,0.023064315,0.0029784844,0.0028763441,0.018565223,0.0012226234],"domain_scores_gemma":[0.9638862,0.014956142,0.0026255632,0.0144455,0.003724421,0.00036217482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026988305,0.0017236358,0.001957113,0.003668871,0.001052536,0.0052879388,0.0033722473,0.0029245273,0.0021352172],"category_scores_gemma":[0.044938736,0.0017054464,0.0034010962,0.0018950395,0.0035464377,0.006877483,0.009351534,0.0040030605,0.0007846476],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002604452,0.00028163768,0.003554268,0.0012108575,0.00057035056,0.0008551225,0.0016710717,0.1944392,0.033671904,0.59618413,0.0015685315,0.16573244],"study_design_scores_gemma":[0.00009548935,0.0003901009,0.0012417721,0.0007409478,0.00032007252,0.0005863977,0.0002984797,0.73635507,0.044079863,0.18374485,0.032022215,0.00012472822],"about_ca_topic_score_codex":0.0018208849,"about_ca_topic_score_gemma":0.0015797007,"teacher_disagreement_score":0.026988305,"about_ca_system_score_codex":0.001833026,"about_ca_system_score_gemma":0.0063545075,"threshold_uncertainty_score":0.14272952},"labels":[],"label_agreement":null},{"id":"W2132705868","doi":"10.5194/gmd-5-1009-2012","title":"Assessing climate model software quality: a defect density analysis of three models","year":2012,"lang":"en","type":"article","venue":"Geoscientific model development","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Executable; Context (archaeology); Software; Climate model; Computer science; Software quality; Quality (philosophy); Trustworthiness; Climate change; Software development; Geography; Geology; Programming language; Computer security","score_opus":0.11359214907245384,"score_gpt":0.32922753876155786,"score_spread":0.215635389689104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132705868","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9959235,0.000036580535,0.0033501415,0.000045823544,0.0000022858105,0.000022527913,0.00014357005,0.00009390382,0.00038169615],"genre_scores_gemma":[0.99734527,0.000020678588,0.0022532973,0.0000052586142,0.0000017460985,0.00001828037,0.00028710434,0.000015780623,0.000052601572],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.997364,0.001147655,0.00022655008,0.00032876412,0.0007764718,0.00015655922],"domain_scores_gemma":[0.89887464,0.07827451,0.007366083,0.007661925,0.0065500117,0.0012728324],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009447434,0.00059847336,0.0005518211,0.0031891055,0.0005392467,0.0014585139,0.0011519613,0.00075160427,0.0006863856],"category_scores_gemma":[0.04079361,0.0003706637,0.0016236975,0.002268321,0.0010890777,0.0015386898,0.0011877912,0.0008626152,0.00009640273],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009321744,0.00087713357,0.6692533,0.00013020764,0.0005898238,0.00024472782,0.0010434145,0.29828146,0.0024169679,0.0033414506,0.001016725,0.021872709],"study_design_scores_gemma":[0.000044898374,0.00044454483,0.18868914,0.000021852482,0.00014732017,0.000114474264,0.0003355581,0.80625564,0.001418156,0.0021641254,0.00031118782,0.000053014177],"about_ca_topic_score_codex":0.014822057,"about_ca_topic_score_gemma":0.010140577,"teacher_disagreement_score":0.99055254,"about_ca_system_score_codex":0.0022626086,"about_ca_system_score_gemma":0.00076501264,"threshold_uncertainty_score":0.049963415},"labels":[],"label_agreement":null},{"id":"W2132887549","doi":"10.1145/1368088.1368114","title":"A comparative analysis of the efficiency of change metrics and static code attributes for defect prediction","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":704,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Eclipse; Data mining; Naive Bayes classifier; Java; Machine learning; Set (abstract data type); Software bug; Artificial intelligence; Decision tree; Code (set theory); Process (computing); Source code; Software; Predictive power; Programming language; Support vector machine","score_opus":0.14473973059311193,"score_gpt":0.32788590260721745,"score_spread":0.18314617201410552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132887549","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.895284,0.0071501904,0.08793506,0.0010519385,0.0000971952,0.00010670662,0.001237242,0.0014945808,0.0056430497],"genre_scores_gemma":[0.9670442,0.001514273,0.028334707,0.00008158648,0.000092093556,0.000038559938,0.001625362,0.00019941252,0.0010697905],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9905725,0.0048088105,0.0005021491,0.0010437708,0.002747515,0.0003252431],"domain_scores_gemma":[0.8043551,0.18098867,0.0026293132,0.0046854964,0.006600782,0.00074064085],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012501372,0.0013854344,0.0011016368,0.0070385127,0.0003532257,0.0013584631,0.0008629843,0.0012139995,0.0009772568],"category_scores_gemma":[0.063380055,0.00035850532,0.00096436794,0.004258174,0.0005178257,0.0030955535,0.00065712706,0.0010297833,0.00072437176],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023993512,0.0007968993,0.18804808,0.00073696574,0.0013510407,0.000255067,0.00035405555,0.10135476,0.00796426,0.0013406073,0.0030376844,0.69236124],"study_design_scores_gemma":[0.00012920988,0.0023328874,0.14903837,0.00020724756,0.0008054527,0.0005648991,0.00031919754,0.82245636,0.018192604,0.0028439325,0.002971368,0.00013859248],"about_ca_topic_score_codex":0.0023477091,"about_ca_topic_score_gemma":0.0026430818,"teacher_disagreement_score":0.98749864,"about_ca_system_score_codex":0.00047403775,"about_ca_system_score_gemma":0.00045113728,"threshold_uncertainty_score":0.066114366},"labels":[],"label_agreement":null},{"id":"W2132992867","doi":"10.1109/icpc.2007.41","title":"Using Bayesian Belief Networks to Predict Change Propagation in Software Systems","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Bayesian network; Program comprehension; Software evolution; Dependency (UML); Java; Probabilistic logic; Software system; Change impact analysis; Software maintenance; Software; Software development; Data mining; Software engineering; Machine learning; Artificial intelligence; Theoretical computer science; Programming language; Software construction","score_opus":0.04411527325135882,"score_gpt":0.2925908469978418,"score_spread":0.248475573746483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132992867","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4736806,0.0004676745,0.52183145,0.00032115067,0.000025116231,0.00010815886,0.00050956017,0.0014929703,0.0015632756],"genre_scores_gemma":[0.9193847,0.00019317947,0.07927424,0.000041148018,0.00002020452,0.00007044031,0.0006026335,0.00006815886,0.00034527495],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99830353,0.0007206673,0.000110782144,0.0002741058,0.00047590837,0.00011499894],"domain_scores_gemma":[0.9619621,0.032165904,0.0027313665,0.00095879927,0.0018839673,0.0002978912],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043634716,0.0011468367,0.0006333777,0.0037851105,0.00048712202,0.0013107936,0.0009398427,0.0014375903,0.0008862372],"category_scores_gemma":[0.03815195,0.0010456219,0.00070921815,0.0015631656,0.0005573789,0.003089454,0.00065877876,0.0014449938,0.00028255244],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003104447,0.00017497705,0.06883892,0.00009593421,0.00018691213,0.00011821241,0.0002864684,0.8492609,0.002219974,0.0023660345,0.00061186816,0.07552931],"study_design_scores_gemma":[0.000006802906,0.000017677618,0.0035519619,0.0000069651023,0.000017395701,0.000016763588,0.000010958135,0.99351925,0.00061348255,0.0021431819,0.000083821695,0.000011804165],"about_ca_topic_score_codex":0.01685964,"about_ca_topic_score_gemma":0.01653928,"teacher_disagreement_score":0.01685964,"about_ca_system_score_codex":0.0012711369,"about_ca_system_score_gemma":0.0007701751,"threshold_uncertainty_score":0.033523023},"labels":[],"label_agreement":null},{"id":"W2132994900","doi":"10.1109/wcre.2011.39","title":"An Entropy Evaluation Approach for Triaging Field Crashes: A Case Study of Mozilla Firefox","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Crash; Triage; Computer science; Entropy (arrow of time); Computer security; Medicine; Operating system","score_opus":0.12250823381080148,"score_gpt":0.3577473723486759,"score_spread":0.2352391385378744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132994900","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90972394,0.00045405922,0.08552764,0.00043712155,0.000032080763,0.00031739596,0.0007695897,0.0008502223,0.0018879898],"genre_scores_gemma":[0.9634602,0.00010780695,0.035301726,0.000028404851,0.00001605341,0.00006205172,0.00054132665,0.00006239691,0.00041993163],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967263,0.0015841041,0.00021911177,0.00037166447,0.00094086997,0.00015796108],"domain_scores_gemma":[0.974002,0.018047893,0.00226191,0.0016257874,0.0034358331,0.00062656257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063089128,0.0007353938,0.00063481514,0.0047488944,0.0005915877,0.0007324119,0.0010224457,0.00070195063,0.0005764999],"category_scores_gemma":[0.024932496,0.00025018438,0.000438357,0.0023070483,0.0005729858,0.0014108407,0.0007698287,0.00068911456,0.00011834715],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016043428,0.0013996717,0.38457716,0.00079709897,0.0004978115,0.0019232698,0.0049192924,0.17585039,0.023416901,0.0061824857,0.0069019073,0.39192972],"study_design_scores_gemma":[0.00006378239,0.0009523694,0.12893727,0.000057113015,0.00013966567,0.0007590989,0.0016337582,0.8503789,0.010614655,0.0041547795,0.0022137912,0.00009484814],"about_ca_topic_score_codex":0.007991619,"about_ca_topic_score_gemma":0.011443341,"teacher_disagreement_score":0.007991619,"about_ca_system_score_codex":0.0011315699,"about_ca_system_score_gemma":0.00062940904,"threshold_uncertainty_score":0.03336513},"labels":[],"label_agreement":null},{"id":"W2133339195","doi":"10.1109/vlhcc.2008.4639063","title":"Towards the next generation of bug tracking systems","year":2008,"lang":"en","type":"article","venue":"Proceedings/Proceedings -- IEEE Symposium on Visual Languages and Human-Centric Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":94,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"sort; Computer science; Eclipse; Card sorting; Software bug; Tracking system; World Wide Web; Tracking (education); Data science; Computer security; Database; Software; Operating system; Engineering","score_opus":0.06921882124549027,"score_gpt":0.3169107127593621,"score_spread":0.24769189151387183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133339195","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09106295,0.024011072,0.71055335,0.1019686,0.0030111363,0.0015013069,0.00158149,0.035599887,0.030710246],"genre_scores_gemma":[0.098751105,0.007850929,0.87399304,0.0057259668,0.0006880284,0.0007013631,0.0032058894,0.0011088104,0.0079749655],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98953956,0.0042252466,0.00093097956,0.0013789561,0.0027547367,0.001170485],"domain_scores_gemma":[0.9551537,0.013088774,0.0031402484,0.0070303855,0.014802119,0.0067847315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04039194,0.00141109,0.0016246785,0.0050093406,0.0013608924,0.007587277,0.0066204933,0.007636711,0.012778254],"category_scores_gemma":[0.05018566,0.001743207,0.0014618929,0.0024857125,0.003001879,0.024692943,0.006661165,0.005181123,0.006422215],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011595614,0.0014093302,0.019347925,0.0017494586,0.00021705359,0.0007415466,0.0034729796,0.0057516107,0.021757223,0.09034833,0.08664205,0.7674028],"study_design_scores_gemma":[0.0007363685,0.0026085528,0.020646105,0.0026320927,0.00045122736,0.0018263566,0.00380785,0.07556885,0.015655244,0.13294354,0.74266344,0.00046032498],"about_ca_topic_score_codex":0.00465724,"about_ca_topic_score_gemma":0.0029870605,"teacher_disagreement_score":0.04039194,"about_ca_system_score_codex":0.0021676486,"about_ca_system_score_gemma":0.0060777185,"threshold_uncertainty_score":0.21361554},"labels":[],"label_agreement":null},{"id":"W2133363731","doi":"10.1145/2000799.2000805","title":"Recommending Adaptive Changes for Framework Evolution","year":2011,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":98,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Code refactoring; Computer science; Software evolution; Software engineering; Open source; Simple (philosophy); Software development; Programming language; Software","score_opus":0.20430693453139503,"score_gpt":0.3413837973744098,"score_spread":0.1370768628430148,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133363731","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46915826,0.003677941,0.46731436,0.0022713763,0.00042967388,0.0016061414,0.0036905827,0.04315315,0.008698518],"genre_scores_gemma":[0.5414737,0.0006407325,0.44720843,0.0004599423,0.00007802441,0.00042538642,0.00589248,0.0014015039,0.0024198098],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99493223,0.0011786463,0.00046117313,0.0016550571,0.0015113349,0.00026158028],"domain_scores_gemma":[0.978534,0.009020343,0.0019846573,0.0045066397,0.005311026,0.00064340804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048439866,0.0017156376,0.0011220205,0.0053145466,0.0010379109,0.0019834077,0.0025727027,0.0022299476,0.0018632421],"category_scores_gemma":[0.038507614,0.0008873825,0.001154219,0.002558654,0.00057140907,0.003448948,0.0013391535,0.002028638,0.00090220204],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056734215,0.00068149855,0.12856078,0.00074021204,0.00033865662,0.000601795,0.0017832118,0.023066405,0.01804238,0.0020023908,0.020609446,0.8030059],"study_design_scores_gemma":[0.00029939387,0.00091199134,0.06358908,0.00044461028,0.0008050622,0.0009759191,0.0021515058,0.82200897,0.035806987,0.0065128156,0.06613196,0.00036173215],"about_ca_topic_score_codex":0.013587486,"about_ca_topic_score_gemma":0.0347345,"teacher_disagreement_score":0.013587486,"about_ca_system_score_codex":0.0012622398,"about_ca_system_score_gemma":0.0023500144,"threshold_uncertainty_score":0.027016759},"labels":[],"label_agreement":null},{"id":"W2133561941","doi":"10.1145/1368088.1368162","title":"Open source software peer review practices","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":235,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Commit; Computer science; Artifact (error); Technical peer review; Open source; Software; Open source software; Software engineering; Construct (python library); Data science; Open-source software development; World Wide Web; Peer review; Database; Operating system; Artificial intelligence","score_opus":0.10945013379595914,"score_gpt":0.3702281267375989,"score_spread":0.2607779929416398,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133561941","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.085031085,0.028991526,0.6356332,0.021502595,0.002528434,0.0061552706,0.0010127561,0.01037271,0.20877244],"genre_scores_gemma":[0.51367253,0.01857061,0.39857158,0.00258443,0.0025951602,0.0038605006,0.0017797321,0.002245742,0.056119703],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.8032332,0.06354702,0.018452395,0.012662687,0.099113286,0.0029913068],"domain_scores_gemma":[0.5273234,0.129784,0.04478854,0.10152145,0.18749247,0.009090049],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.090189435,0.0008189954,0.0010829912,0.0096001085,0.0043971376,0.008568614,0.004263667,0.0027813576,0.0075740293],"category_scores_gemma":[0.27947447,0.001019141,0.00076682307,0.007666829,0.0033726706,0.0081193,0.006816016,0.0023654734,0.008118359],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014813269,0.00032438268,0.02180713,0.0026571627,0.00022119837,0.000747265,0.013920102,0.0031878427,0.009163418,0.05047473,0.04669474,0.85065395],"study_design_scores_gemma":[0.00014723984,0.00048134668,0.026525116,0.0023474453,0.00017338473,0.0039158296,0.0053516263,0.008235304,0.011487321,0.059714,0.8813207,0.00030065066],"about_ca_topic_score_codex":0.0020141273,"about_ca_topic_score_gemma":0.0027330432,"teacher_disagreement_score":0.90981054,"about_ca_system_score_codex":0.0025712424,"about_ca_system_score_gemma":0.011447879,"threshold_uncertainty_score":0.476973},"labels":[],"label_agreement":null},{"id":"W2133750098","doi":"10.1109/promise.2007.4","title":"Complexity Measures for Secure Service-Oriented Software Architectures","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software measurement; Software metric; Software system; Software construction; Software development; Software sizing; Software; Software security assurance; Social software engineering; Software peer review; Software engineering; Computer security; Operating system; Information security; Security service","score_opus":0.03906320280723961,"score_gpt":0.2945958594781885,"score_spread":0.2555326566709489,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133750098","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44625643,0.0015380736,0.53800344,0.0006382502,0.00009787687,0.00038997352,0.0011835697,0.00048313796,0.011409348],"genre_scores_gemma":[0.9043074,0.0002847219,0.09352041,0.000043309636,0.000060234914,0.00040602518,0.0007410095,0.000063412794,0.0005734296],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99401844,0.0015753066,0.0006762918,0.00043208335,0.0030591984,0.00023868965],"domain_scores_gemma":[0.9453897,0.035677753,0.009101308,0.004417467,0.0039170044,0.0014968109],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005242242,0.00069208775,0.00062989903,0.0081646545,0.0008580679,0.0021002456,0.00063794834,0.0008883038,0.0018180868],"category_scores_gemma":[0.049411267,0.00027368407,0.0011067205,0.0040197414,0.0017957261,0.0044456962,0.0017189484,0.0012692871,0.00024079964],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005540805,0.0005460517,0.09789786,0.0009671934,0.0005397705,0.0002506701,0.0025502308,0.2670918,0.016449949,0.40416527,0.004890199,0.2040969],"study_design_scores_gemma":[0.00005289424,0.0009800228,0.09125127,0.00017892814,0.00014889534,0.0006204482,0.0009229182,0.5838862,0.0063743577,0.30654514,0.008878357,0.0001605481],"about_ca_topic_score_codex":0.0011272825,"about_ca_topic_score_gemma":0.0011721484,"teacher_disagreement_score":0.0081646545,"about_ca_system_score_codex":0.0023613316,"about_ca_system_score_gemma":0.00083766645,"threshold_uncertainty_score":0.027723968},"labels":[],"label_agreement":null},{"id":"W2134091519","doi":"10.1109/icsm.2009.5306313","title":"Measuring the progress of projects using the time dependence of code changes","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; KPI-driven code analysis; Code review; Code (set theory); Source code; Software evolution; Software; Source lines of code; Software quality; Track (disk drive); Open source; Software metric; Software engineering; Software development; Database; Programming language; Software construction; Operating system","score_opus":0.07381897592583876,"score_gpt":0.29942623446178196,"score_spread":0.2256072585359432,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134091519","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9029021,0.0010994574,0.08640149,0.00026642391,0.00004243714,0.00015823891,0.004077771,0.0010018351,0.0040502474],"genre_scores_gemma":[0.96591437,0.00046075342,0.02863792,0.000033756685,0.00004174883,0.00015819847,0.003723576,0.00016484509,0.00086475885],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99096245,0.0015326515,0.0011593597,0.0022337718,0.0038121382,0.00029969707],"domain_scores_gemma":[0.8459154,0.07906824,0.044560198,0.009860302,0.017940097,0.002655723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064278105,0.0007320372,0.00066148717,0.0126050515,0.00045165405,0.0019011833,0.0010274774,0.0008771327,0.00080080057],"category_scores_gemma":[0.080185585,0.00058801216,0.00047882413,0.008438111,0.00064274156,0.0039948765,0.0014250604,0.0011643476,0.0006079079],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029305406,0.00017314396,0.81231874,0.00038640073,0.0003572474,0.00020230825,0.0018326336,0.02119316,0.0086958995,0.001695335,0.0013179927,0.15153398],"study_design_scores_gemma":[0.00002179328,0.0005381362,0.8960564,0.00009207465,0.00013829359,0.0005657676,0.0007269163,0.083236806,0.0100349765,0.0035418123,0.0049115894,0.00013540489],"about_ca_topic_score_codex":0.0049949465,"about_ca_topic_score_gemma":0.006410226,"teacher_disagreement_score":0.0126050515,"about_ca_system_score_codex":0.000795794,"about_ca_system_score_gemma":0.0006929269,"threshold_uncertainty_score":0.0339939},"labels":[],"label_agreement":null},{"id":"W2134092469","doi":"10.1002/spe.386","title":"Shimba—an environment for reverse engineering Java software systems","year":2001,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":138,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Academy of Finland; Association of Canadian Universities for Research in Astronomy; Nokia","keywords":"Computer science; Reverse engineering; Java; Sequence diagram; Software; Programming language; Abstraction; TRACE (psycholinguistics); Sequence (biology); Software system; Software engineering; Unified Modeling Language","score_opus":0.02269407837112575,"score_gpt":0.2825363977862558,"score_spread":0.25984231941513003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134092469","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007883008,0.00015389644,0.88193446,0.00010156831,0.00004283955,0.0002494379,0.00036813715,0.10671577,0.002550802],"genre_scores_gemma":[0.054418646,0.0004306491,0.92540425,0.00011105187,0.000025039233,0.0006815664,0.001877555,0.013054864,0.003996366],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979741,0.0007581191,0.00019252308,0.00026737113,0.0006476573,0.00016028664],"domain_scores_gemma":[0.9917589,0.004766831,0.0006190759,0.0018950677,0.00067323924,0.00028704663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004525552,0.0012820527,0.0008012364,0.0013064332,0.0006540711,0.001802105,0.0029454387,0.0010920765,0.012076816],"category_scores_gemma":[0.009567131,0.0019275115,0.0011826989,0.0007235177,0.00080209656,0.004321361,0.0035736046,0.002747303,0.004916313],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001909926,0.0010457085,0.004200649,0.0023559583,0.00023215183,0.0021902495,0.0034361803,0.061883576,0.102001645,0.07965827,0.06101375,0.6800719],"study_design_scores_gemma":[0.0013628653,0.00087878393,0.0032285443,0.00075482554,0.00016602069,0.0019337119,0.0005374445,0.46645328,0.1097339,0.05479305,0.35975024,0.0004072827],"about_ca_topic_score_codex":0.0011626801,"about_ca_topic_score_gemma":0.0015392723,"teacher_disagreement_score":0.012076816,"about_ca_system_score_codex":0.0004689948,"about_ca_system_score_gemma":0.0015187047,"threshold_uncertainty_score":0.040400982},"labels":[],"label_agreement":null},{"id":"W2134347783","doi":"10.1155/2009/829725","title":"Improving Effort Estimation by Voting Software Estimation Models","year":2009,"lang":"en","type":"article","venue":"Advances in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Benchmarking; Estimation; Software; Task (project management); Software metric; Software sizing; Software development; Data mining; Use Case Points; Software project management; Data science; Software engineering; Software quality; Software construction; Systems engineering","score_opus":0.006556375910602392,"score_gpt":0.24842231006292848,"score_spread":0.24186593415232607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134347783","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06436556,0.000365679,0.9322624,0.00025393383,0.000071849725,0.000068632755,0.00018571217,0.0012469196,0.001179346],"genre_scores_gemma":[0.6689037,0.0002629414,0.32729796,0.00014890189,0.00011011907,0.00019779267,0.0014419514,0.00034078557,0.0012958411],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9856198,0.0070232255,0.0013448312,0.0028111462,0.0025022973,0.0006986614],"domain_scores_gemma":[0.9507863,0.030724801,0.0037107444,0.0074238316,0.006831733,0.00052257045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016037047,0.0016986002,0.002993344,0.0038440279,0.0007481331,0.0029477936,0.0034095158,0.002226051,0.0015898818],"category_scores_gemma":[0.07698764,0.001217237,0.0015876761,0.0038082332,0.0005745613,0.0052575297,0.0026842118,0.002312789,0.0009788502],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060099934,0.0003382573,0.041016612,0.00027505748,0.00037597783,0.00013487229,0.00043319666,0.4333005,0.0057405294,0.009991565,0.004360111,0.50343245],"study_design_scores_gemma":[0.000022792716,0.000084011044,0.0023786472,0.0000297473,0.000055383847,0.000042268024,0.000045210778,0.98800313,0.0020087599,0.0064499592,0.0008561766,0.000023941524],"about_ca_topic_score_codex":0.0037039998,"about_ca_topic_score_gemma":0.0050257854,"teacher_disagreement_score":0.016037047,"about_ca_system_score_codex":0.00097015896,"about_ca_system_score_gemma":0.0012374894,"threshold_uncertainty_score":0.084813},"labels":[],"label_agreement":null},{"id":"W2134363520","doi":"10.1109/wcre.2008.57","title":"Integrative Levels of Program Comprehension","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Intension; Comprehension; Computer science; Program comprehension; Percept; Process (computing); Unit (ring theory); Software engineering; Programming language; Software; Mathematics education; Psychology; Software system; Linguistics","score_opus":0.0668548942817134,"score_gpt":0.3212560155372648,"score_spread":0.25440112125555137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134363520","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19896881,0.0012705005,0.6045183,0.0027175248,0.000057896374,0.0006845406,0.00033907214,0.0023349524,0.18910836],"genre_scores_gemma":[0.84553164,0.00042719758,0.14345181,0.00028885753,0.00005076124,0.00048886513,0.00036635087,0.00031371144,0.009080767],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9941275,0.0021045525,0.0004658357,0.0008837686,0.0018596343,0.0005586809],"domain_scores_gemma":[0.98193467,0.00852902,0.0018822287,0.003535905,0.0032425185,0.0008756228],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047201207,0.0008043277,0.0004551354,0.0025162976,0.0013916767,0.006481772,0.0013930278,0.001591701,0.008757541],"category_scores_gemma":[0.022980593,0.00059036486,0.000805629,0.0015103366,0.006474075,0.010019916,0.005698608,0.0032039923,0.0011433515],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016914231,0.00024190794,0.019351473,0.0006205563,0.00008960749,0.0005261613,0.06126248,0.0032529323,0.019250557,0.72354573,0.0026531855,0.16903631],"study_design_scores_gemma":[0.00006475009,0.0006471706,0.040129036,0.0007607131,0.00017524573,0.0017290699,0.019473523,0.025134906,0.018454857,0.8124713,0.080773704,0.0001858662],"about_ca_topic_score_codex":0.000918205,"about_ca_topic_score_gemma":0.0007668286,"teacher_disagreement_score":0.008757541,"about_ca_system_score_codex":0.0016003317,"about_ca_system_score_gemma":0.0021081674,"threshold_uncertainty_score":0.029296875},"labels":[],"label_agreement":null},{"id":"W2134402473","doi":"10.1109/vissoft.2015.7332428","title":"Visualization based API usage patterns refining","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Visualization; Task (project management); Software; Software engineering; Software visualization; World Wide Web; Software development; Component-based software engineering; Data mining; Programming language","score_opus":0.05846739639752362,"score_gpt":0.3194651161984796,"score_spread":0.260997719800956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134402473","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15864986,0.0011996225,0.79828656,0.0011507066,0.00010736637,0.00056087645,0.004378151,0.026170697,0.0094962185],"genre_scores_gemma":[0.42035478,0.0009534059,0.56346613,0.00014601334,0.000044166936,0.0004509018,0.0067230994,0.002530161,0.0053313966],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981692,0.00029629905,0.00022968846,0.00043690545,0.0007020261,0.00016590305],"domain_scores_gemma":[0.99134094,0.0028299822,0.0011504871,0.0016422115,0.0028000637,0.00023630305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015224911,0.0016769001,0.00087318587,0.006858966,0.0007445045,0.0026602473,0.0010990587,0.0009236201,0.0024948854],"category_scores_gemma":[0.012570161,0.00077171304,0.001253498,0.004167063,0.00053337676,0.0033389025,0.0021410713,0.0015224115,0.0010633994],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005815813,0.0003584717,0.08556499,0.0015822427,0.00030434763,0.0018711864,0.009976249,0.016755061,0.07477708,0.022932107,0.023752345,0.7615444],"study_design_scores_gemma":[0.00012352903,0.0003651659,0.079059474,0.0009710289,0.00054752117,0.0045207296,0.004998455,0.55703706,0.12558076,0.067100815,0.15937302,0.000322449],"about_ca_topic_score_codex":0.0062879114,"about_ca_topic_score_gemma":0.009132379,"teacher_disagreement_score":0.006858966,"about_ca_system_score_codex":0.0006050279,"about_ca_system_score_gemma":0.0011630069,"threshold_uncertainty_score":0.012502611},"labels":[],"label_agreement":null},{"id":"W2134852596","doi":"10.1145/1985441.1985468","title":"Software bertillonage","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":99,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Victoria","funders":"","keywords":"Computer science; Java; Source code; Software; Matching (statistics); Software engineering; Software development; Database; World Wide Web; Programming language","score_opus":0.03592069364402037,"score_gpt":0.24003685941897712,"score_spread":0.20411616577495675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134852596","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11307875,0.0029955842,0.5792479,0.0038418786,0.00090610073,0.0006820698,0.0039935233,0.05943283,0.23582143],"genre_scores_gemma":[0.47324872,0.0019619868,0.36789924,0.00069425424,0.00023761047,0.0003478751,0.0074649258,0.0076915855,0.1404537],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99496174,0.0007012934,0.00022717715,0.0007398649,0.0030630755,0.00030688065],"domain_scores_gemma":[0.9893964,0.0025468082,0.00087629486,0.0038315763,0.0027424044,0.00060649146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028934842,0.00070421235,0.00049722864,0.0043230127,0.0016019569,0.003666451,0.0012849376,0.0010705782,0.018549576],"category_scores_gemma":[0.018592944,0.00056670286,0.0006154232,0.0030403764,0.0011407948,0.0053682565,0.0030063016,0.0017610522,0.007422321],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067155936,0.00014440205,0.009778744,0.00033920907,0.000052252144,0.0008100758,0.002561116,0.0048317383,0.014744714,0.16321643,0.052574098,0.7502756],"study_design_scores_gemma":[0.00005395702,0.00031449526,0.0075185494,0.00032465422,0.00004206331,0.0028345496,0.00060957106,0.03543233,0.024207277,0.037865218,0.89065033,0.00014701784],"about_ca_topic_score_codex":0.007269361,"about_ca_topic_score_gemma":0.0074707028,"teacher_disagreement_score":0.018549576,"about_ca_system_score_codex":0.0015801047,"about_ca_system_score_gemma":0.002161351,"threshold_uncertainty_score":0.062054574},"labels":[],"label_agreement":null},{"id":"W2134876531","doi":"10.1109/msr.2007.27","title":"Recommending Emergent Teams","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":128,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Software; Heuristic; Team software process; Software development; World Wide Web; Software engineering; Knowledge management; Data science; Software development process; Artificial intelligence","score_opus":0.020668312632844524,"score_gpt":0.29963938807652296,"score_spread":0.27897107544367844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134876531","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32739404,0.0025003094,0.65413755,0.0012244923,0.00040295493,0.0005821424,0.0024413462,0.002533962,0.008783275],"genre_scores_gemma":[0.7426875,0.0007846736,0.24504972,0.00036697014,0.00035971001,0.00040248822,0.0046146973,0.00023370866,0.0055004936],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99765813,0.0006585441,0.00017895084,0.00080365886,0.0005127876,0.00018783743],"domain_scores_gemma":[0.9882272,0.007257058,0.0008122506,0.0008968793,0.0022066904,0.00059991295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025670326,0.0011712394,0.0013239634,0.008074746,0.0010522658,0.0016617932,0.001717903,0.0021998892,0.0037150385],"category_scores_gemma":[0.01860102,0.0004962211,0.0009863611,0.0024240664,0.00045277763,0.0024696682,0.0013478505,0.0012295786,0.001726197],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010314325,0.00068387896,0.114876635,0.0008776868,0.000915194,0.0008091715,0.001329815,0.09295459,0.0045961584,0.008814985,0.052486617,0.72062385],"study_design_scores_gemma":[0.00021810943,0.00036519108,0.011738522,0.00020917806,0.0004148264,0.00056886853,0.0011155249,0.9472739,0.0052717794,0.020065695,0.012681582,0.000076706856],"about_ca_topic_score_codex":0.0072644143,"about_ca_topic_score_gemma":0.012150414,"teacher_disagreement_score":0.008074746,"about_ca_system_score_codex":0.00079958927,"about_ca_system_score_gemma":0.0013366484,"threshold_uncertainty_score":0.014444232},"labels":[],"label_agreement":null},{"id":"W2135038121","doi":"10.1109/iccbss.2006.19","title":"Maintaining COTS-Based Systems: Start with the Design","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Software engineering; Software maintenance; Software; Systems engineering; Commercial off-the-shelf; Software system; Reliability engineering; Engineering; Operating system","score_opus":0.021649403660443104,"score_gpt":0.23369034358599716,"score_spread":0.21204093992555406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135038121","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09183666,0.019705506,0.7398661,0.04269728,0.001460239,0.00090326485,0.00016862724,0.0018985942,0.10146375],"genre_scores_gemma":[0.30480888,0.009776626,0.62582946,0.0045036064,0.0004941456,0.00043575853,0.0001811272,0.0007366733,0.05323375],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981793,0.000483949,0.00010599227,0.00015335255,0.00096895365,0.00010845966],"domain_scores_gemma":[0.9967777,0.00044895004,0.00024492037,0.0007319869,0.0014487277,0.00034775594],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003188246,0.00059583766,0.0002991297,0.0007125991,0.0012145408,0.0031886066,0.0011720948,0.0015202941,0.002099703],"category_scores_gemma":[0.00523268,0.0005922334,0.00026952566,0.00072525797,0.0023511436,0.0060298773,0.0015250616,0.0019242751,0.0017013726],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007472083,0.00021501926,0.0071303067,0.0021669627,0.00006073967,0.0009860352,0.02016082,0.0030808374,0.036982838,0.17359734,0.041495726,0.71404874],"study_design_scores_gemma":[0.000054713,0.00073119573,0.007941882,0.0014728024,0.00008174027,0.0029103186,0.0072173313,0.0068688076,0.016418876,0.06798463,0.8882238,0.000093873554],"about_ca_topic_score_codex":0.0024633443,"about_ca_topic_score_gemma":0.004774712,"teacher_disagreement_score":0.0031886066,"about_ca_system_score_codex":0.0012769045,"about_ca_system_score_gemma":0.002109225,"threshold_uncertainty_score":0.01686132},"labels":[],"label_agreement":null},{"id":"W2135198476","doi":"10.1145/1368088.1368161","title":"Predicting defects using network analysis on dependency graphs","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":575,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Dependency (UML); Dependency graph; Graph; Face (sociological concept); Software; Code (set theory); Theoretical computer science; Software engineering; Programming language","score_opus":0.028030367664703786,"score_gpt":0.26653112343365976,"score_spread":0.23850075576895596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135198476","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7141137,0.0010439929,0.2727107,0.00065892685,0.00006161141,0.00022113379,0.004503749,0.0024233963,0.0042627663],"genre_scores_gemma":[0.96771073,0.00046289538,0.027811585,0.000035346697,0.000041104755,0.000091995025,0.003073815,0.00007568459,0.0006968625],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990959,0.0003029281,0.000068916925,0.00021864899,0.0002220651,0.00009138102],"domain_scores_gemma":[0.97996175,0.015439501,0.0020312178,0.00067422685,0.0015884071,0.00030485433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015293806,0.001232003,0.0005177026,0.010127763,0.00041564077,0.0008126263,0.0006549965,0.00083092385,0.0014425905],"category_scores_gemma":[0.016952038,0.00041125118,0.0008049047,0.0028912383,0.00034011644,0.0020659685,0.00055940636,0.0006570657,0.00053280813],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040702862,0.00030481475,0.3366243,0.00020397648,0.0004068006,0.0004618295,0.00022512002,0.51625353,0.0026496786,0.003289677,0.0046943547,0.13447894],"study_design_scores_gemma":[0.000007890776,0.000047020683,0.02251022,0.000017358592,0.000055327975,0.00012170118,0.000047520803,0.97130686,0.0007232702,0.0044816495,0.00066548976,0.000015689182],"about_ca_topic_score_codex":0.020316977,"about_ca_topic_score_gemma":0.021202587,"teacher_disagreement_score":0.020316977,"about_ca_system_score_codex":0.0011954387,"about_ca_system_score_gemma":0.00056175684,"threshold_uncertainty_score":0.040397465},"labels":[],"label_agreement":null},{"id":"W2135250953","doi":"10.1109/wpc.2002.1021325","title":"Enhancing program comprehension with recovered state models","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Program comprehension; Computer science; Finite-state machine; Programming language; State (computer science); Transition (genetics); Source code; Representation (politics); Comprehension; Static program analysis; Code (set theory); Class (philosophy); Software; Static analysis; Software engineering; Software system; Theoretical computer science; Artificial intelligence; Software development","score_opus":0.023493397889943073,"score_gpt":0.2602710242265848,"score_spread":0.23677762633664176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135250953","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041765817,0.00007789632,0.9450948,0.00025529068,0.00001457037,0.00009460145,0.00015261894,0.0111477105,0.0013966215],"genre_scores_gemma":[0.3264297,0.00027960807,0.667468,0.00017126706,0.000027763315,0.00019454538,0.00095507794,0.0023170742,0.0021569761],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987143,0.0005339396,0.00007352772,0.000219553,0.0003899343,0.00006871168],"domain_scores_gemma":[0.9876567,0.0085148625,0.00074279244,0.0019831634,0.0010049415,0.00009751804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015492544,0.0013485459,0.0007888926,0.00094449264,0.00025080415,0.0015572343,0.0014276623,0.001314113,0.0037689211],"category_scores_gemma":[0.019097038,0.0005428223,0.0009881585,0.0005908219,0.00071379496,0.0047022942,0.002040704,0.0015196783,0.0008556728],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001114138,0.0008428165,0.0035670663,0.001591186,0.00014558043,0.0012095319,0.005536917,0.18470499,0.23410179,0.045385897,0.007328871,0.5144712],"study_design_scores_gemma":[0.00009615671,0.00025275833,0.0007809814,0.00008576826,0.000107146756,0.00038326654,0.0002986842,0.8475005,0.1088716,0.031949453,0.00960987,0.00006384688],"about_ca_topic_score_codex":0.0007748307,"about_ca_topic_score_gemma":0.0009386108,"teacher_disagreement_score":0.0037689211,"about_ca_system_score_codex":0.00039989603,"about_ca_system_score_gemma":0.00055084703,"threshold_uncertainty_score":0.01260829},"labels":[],"label_agreement":null},{"id":"W2135303064","doi":"10.1109/wcre.2012.32","title":"An Empirical Study on Factors Impacting Bug Fixing Time","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Universität Zürich; Technische Universiteit Delft","keywords":"Software regression; Software bug; Computer science; Process (computing); Security bug; Empirical research; Code (set theory); Open source; Software maintenance; Software engineering; Software; Software development; Software quality; Programming language; Operating system; Software security assurance; Statistics; Cloud computing","score_opus":0.05837254724366758,"score_gpt":0.3754761812513734,"score_spread":0.3171036340077058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135303064","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99830973,0.00012445223,0.00042385346,0.00009410206,0.00000408859,0.000041277857,0.000153397,0.000013648792,0.0008354551],"genre_scores_gemma":[0.9986958,0.00012753622,0.00064949464,0.000018687553,0.000005804309,0.000042572705,0.00021441931,0.000013752989,0.0002320111],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99067044,0.004066738,0.0011911178,0.0009423721,0.0022957346,0.00083359214],"domain_scores_gemma":[0.54558176,0.35497412,0.06920626,0.007481871,0.016576415,0.0061795716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0108385375,0.0005249107,0.00037221966,0.0034556005,0.0008248707,0.0019130792,0.0010881905,0.00085861783,0.003268768],"category_scores_gemma":[0.15484251,0.00046768077,0.0006682784,0.005145307,0.001049521,0.002319184,0.0007862981,0.0015287371,0.00058881025],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010939162,0.00027201086,0.99077386,0.000055186945,0.000046825655,0.000110966575,0.0013875604,0.00028894973,0.0002000624,0.00013052084,0.00014925622,0.0064753024],"study_design_scores_gemma":[0.000013501598,0.00029079596,0.99421024,0.000032965596,0.000044373086,0.00016815268,0.0028424556,0.0013032627,0.00025693382,0.00012416631,0.0006955826,0.000017589547],"about_ca_topic_score_codex":0.0064719925,"about_ca_topic_score_gemma":0.0057149907,"teacher_disagreement_score":0.0108385375,"about_ca_system_score_codex":0.0015812094,"about_ca_system_score_gemma":0.002161357,"threshold_uncertainty_score":0.057320356},"labels":[],"label_agreement":null},{"id":"W2135370021","doi":"10.1109/iri.2007.4296695","title":"A practical method for the software fault-prediction","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Data mining; Software; Software bug; Fault (geology); Representation (politics); Machine learning; Software metric; Fuzzy logic; Software fault tolerance; Artificial intelligence; Software quality; Software development; Programming language","score_opus":0.04919603470099245,"score_gpt":0.38521949984585535,"score_spread":0.3360234651448629,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135370021","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010952046,0.000067996414,0.99784744,0.000055447854,0.000041022828,0.00003395574,0.000035540575,0.00057820376,0.00024520894],"genre_scores_gemma":[0.10161152,0.0001802916,0.8954479,0.000108293294,0.000118309625,0.00025654727,0.0003221694,0.00010986835,0.0018451007],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99836546,0.000369463,0.00007984295,0.0003931012,0.00071738,0.00007471651],"domain_scores_gemma":[0.99832636,0.00065530115,0.00014388598,0.00023063827,0.00058522826,0.000058583886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014742496,0.0009255312,0.00091362774,0.0021224804,0.00067091803,0.00065444084,0.0012700586,0.0013973602,0.0036563117],"category_scores_gemma":[0.0054480187,0.00036576274,0.0005790976,0.0016172758,0.00066398043,0.0010579015,0.00093609997,0.0015747172,0.0017476403],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013560441,0.00013668295,0.0015481127,0.00020810122,0.000061297054,0.00008342244,0.00005903634,0.11418199,0.013040035,0.010775197,0.0074568028,0.85231364],"study_design_scores_gemma":[0.000014881016,0.000039919214,0.0004752024,0.00001549219,0.000007833193,0.00014039832,0.000013872616,0.9841988,0.0038710134,0.008016388,0.0031884569,0.000017695658],"about_ca_topic_score_codex":0.0019804433,"about_ca_topic_score_gemma":0.002224373,"teacher_disagreement_score":0.0036563117,"about_ca_system_score_codex":0.000584263,"about_ca_system_score_gemma":0.0011675173,"threshold_uncertainty_score":0.012231529},"labels":[],"label_agreement":null},{"id":"W2135424716","doi":"10.1109/chase.2009.5071421","title":"Challenges in the user interface design of an IDE tool recommender","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Recommender system; Focus (optics); User interface; Interface (matter); Software; World Wide Web; Software engineering; Human–computer interaction; Work (physics); Engineering; Programming language","score_opus":0.12491254661772097,"score_gpt":0.3347208093071852,"score_spread":0.2098082626894642,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135424716","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05205055,0.0017487564,0.9250737,0.0060508326,0.00019508535,0.00082836347,0.0001659074,0.0034804402,0.010406428],"genre_scores_gemma":[0.17099267,0.000993292,0.81874347,0.0006900644,0.00008679611,0.00069114054,0.0002542506,0.0005650026,0.0069833854],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98550665,0.008070948,0.0013697803,0.0015289704,0.0030591353,0.000464532],"domain_scores_gemma":[0.97221845,0.0132202385,0.00069822965,0.0032553964,0.009862617,0.0007450237],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023680348,0.00093408115,0.0018720992,0.0015351387,0.0015993213,0.0075969025,0.0037460467,0.0048363796,0.0050537484],"category_scores_gemma":[0.065772735,0.001783337,0.0014674957,0.001169175,0.0015499073,0.009466137,0.0017735889,0.0033586686,0.004990875],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008517079,0.0008184432,0.018902754,0.0018610199,0.00040795913,0.0015107269,0.012531289,0.018947406,0.041037038,0.063076064,0.023872752,0.8161829],"study_design_scores_gemma":[0.0007061064,0.0026796116,0.014944041,0.0009526176,0.00076139235,0.005377036,0.008544855,0.6964916,0.030488702,0.05088104,0.18749048,0.0006825265],"about_ca_topic_score_codex":0.007671775,"about_ca_topic_score_gemma":0.006287235,"teacher_disagreement_score":0.023680348,"about_ca_system_score_codex":0.00092639273,"about_ca_system_score_gemma":0.0014678773,"threshold_uncertainty_score":0.12523514},"labels":[],"label_agreement":null},{"id":"W2135672482","doi":"10.1109/icpc.2011.45","title":"Conflict-Aware Optimal Scheduling of Code Clone Refactoring: A Constraint Programming Approach","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Maintainability; Constraint programming; Programming language; Schedule; Code (set theory); Scheduling (production processes); Software engineering; Set (abstract data type); Software; Engineering; Operating system","score_opus":0.09139369379814016,"score_gpt":0.2899794458955976,"score_spread":0.19858575209745744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135672482","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010604105,0.00030935212,0.9845889,0.0003781513,0.000046612404,0.00016797236,0.00015110862,0.0001556538,0.003598158],"genre_scores_gemma":[0.20341024,0.0007274849,0.79119456,0.00022885477,0.00008116427,0.00063407683,0.00041284275,0.00021700436,0.0030937498],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982235,0.00074863463,0.000079100115,0.00028245887,0.0004139889,0.0002522649],"domain_scores_gemma":[0.9953555,0.0033945595,0.0004007978,0.00013909445,0.0005211306,0.00018898002],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028623,0.0016655315,0.0018004505,0.0014300027,0.000802976,0.0018977238,0.0029857873,0.0014015384,0.0034053188],"category_scores_gemma":[0.006979912,0.0013737038,0.0015697611,0.0025215335,0.0010372067,0.0016634642,0.0011338519,0.002072115,0.00031843822],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035823345,0.000052851756,0.00023416185,0.00008442973,0.00003362007,0.00006172895,0.00005253131,0.97269374,0.0006410643,0.012446838,0.0007867995,0.012876502],"study_design_scores_gemma":[0.000013501805,0.000013235014,0.000052856012,0.000008047294,0.000009388916,0.0000078902,0.000013239634,0.99518365,0.00021986698,0.0040299455,0.00044273655,0.000005749162],"about_ca_topic_score_codex":0.022356067,"about_ca_topic_score_gemma":0.01725686,"teacher_disagreement_score":0.022356067,"about_ca_system_score_codex":0.002377782,"about_ca_system_score_gemma":0.004015807,"threshold_uncertainty_score":0.044451892},"labels":[],"label_agreement":null},{"id":"W2135810082","doi":"10.1109/icsm.2005.95","title":"Towards experience-based mentoring of evolutionary development","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Sequence diagram; Unified Modeling Language; Software evolution; Class diagram; Software engineering; Object-oriented design; Object-oriented programming; Sequence (biology); Software development; Software system; Software design; Set (abstract data type); Heuristic; Software; Programming language; Artificial intelligence; Software construction","score_opus":0.02821801675591105,"score_gpt":0.2820500537264698,"score_spread":0.2538320369705587,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135810082","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08636619,0.0008018989,0.82777435,0.013426732,0.00015314217,0.00028799163,0.000040547806,0.0010330874,0.07011601],"genre_scores_gemma":[0.5294509,0.00066395313,0.45718375,0.00081303244,0.00010352786,0.00028485194,0.00010149109,0.00017958689,0.011218927],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9902243,0.006923032,0.00022610645,0.000816484,0.0013823159,0.00042781048],"domain_scores_gemma":[0.9760978,0.011154989,0.001850079,0.005756305,0.0028997194,0.0022410501],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010095237,0.00083805283,0.00037008902,0.0009914336,0.0011693905,0.004542173,0.0031100698,0.0025453896,0.0044802916],"category_scores_gemma":[0.03399871,0.00041731854,0.0005312172,0.00059623417,0.0038550044,0.0059056887,0.0072044753,0.002879684,0.0010032859],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014826498,0.0013212761,0.010764294,0.00056025264,0.00008681844,0.001030112,0.047908794,0.013925618,0.008440018,0.25420803,0.0075694006,0.6540371],"study_design_scores_gemma":[0.0002261176,0.0010380659,0.0075511215,0.00076114817,0.000113668495,0.0017378654,0.020321896,0.120552994,0.018105399,0.5347072,0.2946829,0.0002017417],"about_ca_topic_score_codex":0.0010177449,"about_ca_topic_score_gemma":0.0015371309,"teacher_disagreement_score":0.010095237,"about_ca_system_score_codex":0.0014547387,"about_ca_system_score_gemma":0.0023106774,"threshold_uncertainty_score":0.05338937},"labels":[],"label_agreement":null},{"id":"W2135875948","doi":"10.1007/s00766-003-0181-1","title":"User?s manual as a requirements specification: case studies","year":2004,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software requirements specification; Computer science; Requirements analysis; User requirements document; Software engineering; Programming language; Software; Software development; Software design","score_opus":0.06572052522864251,"score_gpt":0.35026638607864796,"score_spread":0.28454586085000544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135875948","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91964823,0.00035305493,0.025791781,0.0018053644,0.000056845078,0.0008479939,0.0004919595,0.0004702793,0.05053441],"genre_scores_gemma":[0.946822,0.0002687751,0.03444545,0.00036349494,0.000021749898,0.0004460362,0.0005591824,0.00025431515,0.016818965],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9869718,0.009168634,0.00049282616,0.0005210825,0.0024226764,0.00042296087],"domain_scores_gemma":[0.9407956,0.0484691,0.0019079145,0.0047188476,0.0029752175,0.0011332672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00822201,0.0004921638,0.00052720646,0.0014913881,0.002921195,0.0024062137,0.0024654933,0.0043481914,0.0076127728],"category_scores_gemma":[0.027463952,0.0005286983,0.00056259433,0.0018329831,0.0016615749,0.0021005163,0.0019622664,0.0015280407,0.0021396321],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039199223,0.0187914,0.038823176,0.003017463,0.00015503963,0.05730956,0.12413944,0.037728667,0.036234237,0.06538522,0.047885846,0.56661004],"study_design_scores_gemma":[0.002123919,0.014111194,0.06404684,0.0019442727,0.0004606293,0.043507863,0.12296986,0.21060878,0.08147895,0.027159283,0.43080574,0.00078275765],"about_ca_topic_score_codex":0.006314544,"about_ca_topic_score_gemma":0.012566019,"teacher_disagreement_score":0.00822201,"about_ca_system_score_codex":0.002316878,"about_ca_system_score_gemma":0.0018088605,"threshold_uncertainty_score":0.04348266},"labels":[],"label_agreement":null},{"id":"W2135937426","doi":"10.1109/is.2008.4670508","title":"A Bayesian approach for software quality prediction","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software quality; Gibbs sampling; Data mining; Bayesian probability; Software; Software metric; Machine learning; Dirichlet distribution; Software development; Artificial intelligence; Mathematics; Programming language","score_opus":0.05145918430503104,"score_gpt":0.2931873281437155,"score_spread":0.24172814383868446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135937426","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001636833,0.00022593702,0.9972609,0.0001278285,0.000015230871,0.000014365861,0.000040239483,0.0001571667,0.0005214726],"genre_scores_gemma":[0.23849875,0.0013573222,0.75429964,0.00030503032,0.00032687595,0.00035811152,0.00058688765,0.0001889741,0.0040783635],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964121,0.0013716085,0.00015562508,0.0006767664,0.0012043479,0.00017964155],"domain_scores_gemma":[0.9931914,0.005046134,0.000361949,0.0004504309,0.0007981311,0.00015200078],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042127683,0.0009903498,0.0016569147,0.0032832357,0.0007914621,0.0017599101,0.0029787645,0.002040543,0.0025153537],"category_scores_gemma":[0.017267415,0.0010743226,0.0013557103,0.0024445853,0.0016074842,0.003232435,0.0019083332,0.0023162034,0.00092219154],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013336171,0.00014163203,0.0026144483,0.0001456907,0.00020069312,0.00012240701,0.0002079699,0.5825032,0.0021220054,0.14922117,0.0036717802,0.25891572],"study_design_scores_gemma":[0.0000137690195,0.000020059348,0.00033868538,0.000017533153,0.000019869543,0.00005317842,0.000008548441,0.91051656,0.0003511266,0.08713836,0.0014965977,0.000025744801],"about_ca_topic_score_codex":0.0074873813,"about_ca_topic_score_gemma":0.0062575317,"teacher_disagreement_score":0.0074873813,"about_ca_system_score_codex":0.0015967018,"about_ca_system_score_gemma":0.0015295402,"threshold_uncertainty_score":0.022279501},"labels":[],"label_agreement":null},{"id":"W2135990152","doi":"","title":"The Effects of Business Process Representation Type on the Assessment of Business and Control Risk: Diagrams Versus Narratives","year":2011,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Diagrammatic reasoning; Representation (politics); Narrative; Control (management); Affect (linguistics); Process (computing); Task (project management); Perception; Computer science; Psychology; Artificial intelligence; Linguistics; Political science; Management","score_opus":0.015184647058174556,"score_gpt":0.285766632583953,"score_spread":0.27058198552577845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135990152","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9927539,0.00018785209,0.0022658184,0.0001746846,0.000026102687,0.00012069437,0.000036342364,0.000042247957,0.0043923133],"genre_scores_gemma":[0.9940399,0.00019658275,0.0041114683,0.00008087823,0.000026599762,0.0001652552,0.00007888491,0.000040334256,0.001260263],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9819916,0.0121994335,0.0017601922,0.0012259608,0.002411745,0.00041112292],"domain_scores_gemma":[0.5295093,0.42400363,0.029892111,0.007915639,0.004874459,0.0038048914],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013473467,0.0007302782,0.00061389914,0.0013154022,0.0005502129,0.004874298,0.0009340673,0.0012864561,0.0065285303],"category_scores_gemma":[0.21218643,0.00053168647,0.00085362827,0.00086370826,0.0017739572,0.0054713343,0.0020973014,0.0017389859,0.00072199747],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.049431436,0.023565765,0.40503493,0.002835307,0.0015246677,0.000826026,0.0505655,0.0135602495,0.07903244,0.014608421,0.0015230416,0.35749227],"study_design_scores_gemma":[0.0026336908,0.032154188,0.83831304,0.0013660594,0.0018737043,0.0009696386,0.020535707,0.03372888,0.04054521,0.01900531,0.008314354,0.0005602822],"about_ca_topic_score_codex":0.00073230756,"about_ca_topic_score_gemma":0.0006789064,"teacher_disagreement_score":0.013473467,"about_ca_system_score_codex":0.0006824005,"about_ca_system_score_gemma":0.0005353609,"threshold_uncertainty_score":0.071255326},"labels":[],"label_agreement":null},{"id":"W2136112704","doi":"10.1109/wpc.2004.1311060","title":"Using development history sticky notes to understand software architecture","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Program slicing; Software engineering; Software system; Source code; Software evolution; Software development; Backporting; Dependency graph; Software construction; Software architecture; Resource-oriented architecture; Software analytics; Software; Operating system","score_opus":0.07042513991851523,"score_gpt":0.27666314089034105,"score_spread":0.20623800097182582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136112704","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062301517,0.0013314462,0.90599334,0.0026789142,0.0001243119,0.00017268564,0.0007220215,0.0034428083,0.023232954],"genre_scores_gemma":[0.26301298,0.0024899642,0.72299236,0.0002016445,0.0000643623,0.00012818686,0.0008374544,0.00062209816,0.009650903],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9996257,0.000116080846,0.000042841508,0.000074471274,0.00011479413,0.000026043412],"domain_scores_gemma":[0.99534565,0.0023770959,0.000694274,0.0010106001,0.000436094,0.00013629928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010386404,0.00069935416,0.00019070403,0.0031186412,0.0008623686,0.002057996,0.0008909357,0.00088131643,0.005182829],"category_scores_gemma":[0.01008892,0.0005789679,0.00034620427,0.0014684057,0.001898478,0.009061513,0.0015329293,0.0019414066,0.00089156796],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012705303,0.000069261645,0.007333192,0.00046894755,0.00002826449,0.0017757423,0.021039322,0.025284182,0.009668133,0.3740676,0.008750242,0.551388],"study_design_scores_gemma":[0.000049291833,0.00012093517,0.010352932,0.0008544367,0.00006132187,0.0016442736,0.007777024,0.12425286,0.013419963,0.60342795,0.23788042,0.00015859243],"about_ca_topic_score_codex":0.007881933,"about_ca_topic_score_gemma":0.0109764505,"teacher_disagreement_score":0.007881933,"about_ca_system_score_codex":0.0010232204,"about_ca_system_score_gemma":0.001366304,"threshold_uncertainty_score":0.017338276},"labels":[],"label_agreement":null},{"id":"W2136128399","doi":"10.1109/wcre.2008.54","title":"An Empirical Study of Function Clones in Open Source Software","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":137,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Cloning (programming); Java; Source code; Computer science; Open source; Benchmark (surveying); Software maintenance; Function (biology); Linux kernel; Kernel (algebra); Identification (biology); Software system; Software; Operating system; Programming language; Biology; Genetics; Gene; Mathematics","score_opus":0.05894628892791417,"score_gpt":0.34110329098145586,"score_spread":0.2821570020535417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136128399","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9977756,0.00013107383,0.001528082,0.00003323614,0.000001189903,0.00002142339,0.00011669828,0.000018089884,0.0003745487],"genre_scores_gemma":[0.99699104,0.000088981455,0.0022638352,0.000019718449,0.0000036252338,0.00003107768,0.00038559627,0.000015268615,0.00020071068],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9888946,0.004068847,0.0010246953,0.0016693834,0.0039581927,0.00038424094],"domain_scores_gemma":[0.7726737,0.16366938,0.029688329,0.012861894,0.019307153,0.0017995146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066597387,0.00030409475,0.00030063585,0.0036746836,0.0006922355,0.00111179,0.0008423084,0.0006962426,0.0007485802],"category_scores_gemma":[0.08792744,0.00029204253,0.0002862978,0.0032987439,0.0016794852,0.0023734006,0.0014686743,0.00085552677,0.0002154324],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012153251,0.00019281745,0.96532995,0.00012084563,0.0000675113,0.00030447854,0.0045644743,0.0006884688,0.0022676308,0.0005064374,0.0003120399,0.025523836],"study_design_scores_gemma":[0.000017117709,0.00050691847,0.9755636,0.00007569172,0.000061397885,0.0021427525,0.004900076,0.008270704,0.0052981065,0.00047553054,0.002651147,0.00003704831],"about_ca_topic_score_codex":0.0021463556,"about_ca_topic_score_gemma":0.0031408824,"teacher_disagreement_score":0.0066597387,"about_ca_system_score_codex":0.0007346639,"about_ca_system_score_gemma":0.0005461895,"threshold_uncertainty_score":0.035220444},"labels":[],"label_agreement":null},{"id":"W2136183234","doi":"10.1109/icpc.2008.37","title":"Mendel: A Model, Metrics, and Rules to Understand Class Hierarchies","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Inheritance (genetic algorithm); Computer science; Code reuse; Programming language; Programmer; Multiple inheritance; Class (philosophy); GRASP; Object-oriented programming; Class hierarchy; Subtyping; Theoretical computer science; Reuse; Artificial intelligence; Software","score_opus":0.05435968334041315,"score_gpt":0.27126145544828245,"score_spread":0.21690177210786932,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136183234","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003314546,0.00036093584,0.98598146,0.0010916379,0.00005139996,0.00017533755,0.0008445437,0.004631985,0.0035481923],"genre_scores_gemma":[0.017264845,0.00027242242,0.979037,0.00014392707,0.000025666031,0.00026763836,0.00072161114,0.0005562712,0.0017106488],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.993899,0.0012358878,0.0009124139,0.00085305324,0.0028754715,0.00022416061],"domain_scores_gemma":[0.98644096,0.0060957265,0.0024883728,0.0020365133,0.002397282,0.00054119725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070454376,0.0017239404,0.0011389663,0.0071400907,0.0014711462,0.0040956284,0.004093027,0.0019338746,0.002943742],"category_scores_gemma":[0.030014703,0.0013091224,0.0014982997,0.003385765,0.0021603608,0.008528873,0.002543705,0.0023428772,0.0011241934],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013736331,0.00027204357,0.009070308,0.0009971922,0.00013764761,0.0007915594,0.0017969472,0.063820586,0.0042956998,0.57569396,0.024087008,0.31889975],"study_design_scores_gemma":[0.00010656224,0.00012249961,0.0022705556,0.0006147663,0.00013775418,0.00095636037,0.000261953,0.29847038,0.0065885894,0.4997489,0.19051932,0.00020245067],"about_ca_topic_score_codex":0.016365916,"about_ca_topic_score_gemma":0.029023485,"teacher_disagreement_score":0.016365916,"about_ca_system_score_codex":0.0026050059,"about_ca_system_score_gemma":0.0045976965,"threshold_uncertainty_score":0.037260234},"labels":[],"label_agreement":null},{"id":"W2136224168","doi":"10.1109/wcre.2004.21","title":"Fingerprinting design patterns","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":150,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Architecture; Motif (music); Software; Software design pattern; Identification (biology); Theoretical computer science; Metric (unit); Artificial intelligence; Data mining; Software engineering; Machine learning; Programming language; Engineering","score_opus":0.03133773792809324,"score_gpt":0.26696091154221974,"score_spread":0.2356231736141265,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136224168","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23072964,0.0016116359,0.73673475,0.00068295456,0.00013212507,0.00056492386,0.0050284225,0.006679608,0.01783591],"genre_scores_gemma":[0.4587906,0.00064091384,0.52727306,0.00017894169,0.000027986836,0.00042207327,0.004350608,0.0011587335,0.00715707],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99687564,0.0006639172,0.00026810804,0.00070616655,0.0013099714,0.00017622909],"domain_scores_gemma":[0.9884112,0.0057533933,0.0011251944,0.0030203308,0.0014615452,0.00022848528],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00153921,0.001058511,0.0008783705,0.003169463,0.000741324,0.0017062033,0.0014487058,0.0010506616,0.0057693203],"category_scores_gemma":[0.016955502,0.00066413963,0.00067971624,0.003259841,0.00101873,0.003617584,0.0016893953,0.0010384925,0.0012937081],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006001094,0.00033210064,0.02323742,0.0013043679,0.00012951285,0.0005952161,0.0016159483,0.029374208,0.02667497,0.05761129,0.012131095,0.8463937],"study_design_scores_gemma":[0.00027573027,0.00096453447,0.017617924,0.00056786614,0.00030459324,0.005091536,0.0022104403,0.4483091,0.10365505,0.23125999,0.18949184,0.00025141131],"about_ca_topic_score_codex":0.00152177,"about_ca_topic_score_gemma":0.0028708102,"teacher_disagreement_score":0.0057693203,"about_ca_system_score_codex":0.00083169783,"about_ca_system_score_gemma":0.0010874469,"threshold_uncertainty_score":0.019300282},"labels":[],"label_agreement":null},{"id":"W2136366285","doi":"10.1109/iwse.2005.1","title":"A Light-Weight Proactive Software Change Impact Analysis Using Use Case Maps","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Change impact analysis; Computer science; Software engineering; Software; Abstraction; Focus (optics); Business requirements; Software development; Interface (matter); Business process; Engineering; Programming language; Operating system; Work in process","score_opus":0.048472450201303474,"score_gpt":0.30031307310382765,"score_spread":0.2518406229025242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136366285","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27652323,0.00025379827,0.70529586,0.0004519537,0.00002480349,0.0011348039,0.0013776077,0.0058432855,0.009094622],"genre_scores_gemma":[0.6083543,0.00018916545,0.38819245,0.000044325574,0.000015144085,0.00045632067,0.0011865685,0.00017220918,0.0013896046],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9944654,0.0018231238,0.0002936536,0.00038692917,0.002836773,0.00019418148],"domain_scores_gemma":[0.98461574,0.008969663,0.0011531203,0.002569464,0.0024554688,0.00023651481],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031010865,0.0014181748,0.0007338763,0.011267933,0.0010401114,0.0032078244,0.0012768,0.0010595587,0.0024957233],"category_scores_gemma":[0.018779844,0.00074990507,0.0013695569,0.0037772881,0.0006349113,0.003307348,0.0016648188,0.0008297798,0.00060923386],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005756268,0.0010765115,0.09005622,0.00073067745,0.0006008752,0.0015399278,0.0028700565,0.069964916,0.04203851,0.015408746,0.0038990185,0.77123904],"study_design_scores_gemma":[0.000084834166,0.0007039422,0.070048615,0.00025056413,0.00042461892,0.001618759,0.002187754,0.844031,0.04277992,0.02351432,0.014065035,0.00029066304],"about_ca_topic_score_codex":0.006597246,"about_ca_topic_score_gemma":0.00586652,"teacher_disagreement_score":0.011267933,"about_ca_system_score_codex":0.0008622216,"about_ca_system_score_gemma":0.0013466239,"threshold_uncertainty_score":0.016400278},"labels":[],"label_agreement":null},{"id":"W2136695057","doi":"10.1109/wcre.2012.18","title":"Understanding Android Fragmentation with Topic Analysis of Vendor-Specific Bugs","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":116,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Software portability; Computer science; Vendor; Android (operating system); Fragmentation (computing); Software bug; World Wide Web; Operating system; Software","score_opus":0.09287072505236618,"score_gpt":0.2805145327793649,"score_spread":0.1876438077269987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136695057","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86107624,0.0015001065,0.13236485,0.0005764314,0.00006540222,0.00016736789,0.0018581585,0.0006555777,0.0017358984],"genre_scores_gemma":[0.9589207,0.00040245897,0.03701365,0.000049708968,0.00008677984,0.00017676415,0.0026676646,0.00009207383,0.00059017807],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9979005,0.0008066286,0.00020936696,0.00054953684,0.00031477786,0.00021923563],"domain_scores_gemma":[0.97405857,0.016120141,0.004861839,0.0018633184,0.0025716121,0.0005244655],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047514304,0.00070847094,0.000750863,0.007923565,0.0008209023,0.0019556326,0.0007172098,0.00079581485,0.0005730489],"category_scores_gemma":[0.019532016,0.00049805635,0.0011429192,0.0045607807,0.0006285842,0.0029640896,0.0014413293,0.0009934953,0.00039351688],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064754515,0.00036656277,0.6619374,0.00062244275,0.00034874378,0.0006934553,0.013041481,0.022624642,0.014999976,0.004463256,0.00690622,0.2733482],"study_design_scores_gemma":[0.00007056331,0.00027774007,0.5342725,0.00015149942,0.00031905615,0.0011036481,0.006071787,0.4294314,0.005789122,0.013403084,0.008966098,0.00014358458],"about_ca_topic_score_codex":0.009694375,"about_ca_topic_score_gemma":0.009258303,"teacher_disagreement_score":0.009694375,"about_ca_system_score_codex":0.00090504804,"about_ca_system_score_gemma":0.00083875266,"threshold_uncertainty_score":0.025128245},"labels":[],"label_agreement":null},{"id":"W2136848333","doi":"10.1109/ccece.2005.1557152","title":"Intelligent software measurement system (ISMS)","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Software measurement; Software engineering; Software system; Software; Software construction; Software metric; Software development; Programming language","score_opus":0.02911181045085277,"score_gpt":0.23649716696886708,"score_spread":0.20738535651801432,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136848333","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019773703,0.00085381494,0.83945596,0.00080687075,0.00019979262,0.0006058821,0.0015828436,0.09316979,0.043551363],"genre_scores_gemma":[0.36780927,0.0010026162,0.60255456,0.0006078925,0.00015830187,0.0008314224,0.0063457834,0.0015897678,0.019100491],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99661726,0.0006291423,0.00036228562,0.00060731976,0.0016167477,0.0001673578],"domain_scores_gemma":[0.9957651,0.0009472599,0.0006536944,0.0011271584,0.0012031688,0.0003036322],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029170744,0.00062893564,0.0009851534,0.0023630508,0.0007706603,0.0028717031,0.0015763155,0.0011180771,0.006167048],"category_scores_gemma":[0.0076105106,0.00044317008,0.00052197004,0.0017123789,0.00080479187,0.0029595064,0.002145347,0.0013048353,0.0043780445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005911833,0.00043057255,0.007502754,0.00082726305,0.00020000924,0.00030825735,0.0010446208,0.019899018,0.029324573,0.10431118,0.065986425,0.76957417],"study_design_scores_gemma":[0.00028101946,0.00081551255,0.0102203265,0.00044824864,0.00036257377,0.0008391135,0.0002768882,0.3345813,0.0727802,0.05948286,0.51968586,0.00022614846],"about_ca_topic_score_codex":0.0020470815,"about_ca_topic_score_gemma":0.0014371133,"teacher_disagreement_score":0.006167048,"about_ca_system_score_codex":0.0014603267,"about_ca_system_score_gemma":0.002700262,"threshold_uncertainty_score":0.020630836},"labels":[],"label_agreement":null},{"id":"W2136988678","doi":"10.1109/icpc.2008.10","title":"Reusing Program Investigation Knowledge for Code Understanding","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Reuse; Computer science; Task (project management); Set (abstract data type); Software engineering; Software; Software maintenance; Code (set theory); Code reuse; Software development; Programming language; Systems engineering; Engineering","score_opus":0.2066191496039817,"score_gpt":0.34864216034695084,"score_spread":0.14202301074296914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136988678","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19668955,0.0019070794,0.78463286,0.0018445096,0.000035920773,0.00057454495,0.00055258034,0.0065581594,0.0072048726],"genre_scores_gemma":[0.5662198,0.00084154226,0.42938977,0.00023909047,0.00004241044,0.0001707922,0.001209033,0.0007484225,0.0011392082],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9908236,0.0035058379,0.0007798351,0.0016694446,0.002772209,0.00044912443],"domain_scores_gemma":[0.8765933,0.07402287,0.007346659,0.03393842,0.0073693944,0.00072936865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010262508,0.0016205559,0.0015796039,0.009598734,0.0015109975,0.0049872813,0.0033384487,0.0022615227,0.0022262386],"category_scores_gemma":[0.119800426,0.0015619469,0.0018526,0.004403931,0.0022128767,0.01736546,0.004742883,0.0030185017,0.0007329998],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027283962,0.00068775244,0.026250362,0.0009474944,0.00023286077,0.000645321,0.01176259,0.013926928,0.015267785,0.009749643,0.0026472998,0.9176092],"study_design_scores_gemma":[0.0003172981,0.0017319791,0.079331174,0.0025849864,0.002168706,0.004219858,0.010600025,0.5545152,0.10562465,0.15456031,0.083503686,0.00084206625],"about_ca_topic_score_codex":0.010069475,"about_ca_topic_score_gemma":0.014567819,"teacher_disagreement_score":0.010262508,"about_ca_system_score_codex":0.0017415804,"about_ca_system_score_gemma":0.0048162597,"threshold_uncertainty_score":0.054273963},"labels":[],"label_agreement":null},{"id":"W2137711347","doi":"10.1109/icsm.2015.7332481","title":"PARC: Recommending API methods parameters","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Programming language; Task (project management); Code (set theory); Plug-in; Eclipse; Code generation; Source code; Software engineering; Unreachable code; Redundant code; Operating system; Set (abstract data type); Engineering","score_opus":0.14295413883641483,"score_gpt":0.40401840912723247,"score_spread":0.26106427029081763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2137711347","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027267497,0.0020490559,0.76215947,0.0009373437,0.00034623654,0.0009370559,0.010670086,0.18556978,0.010063517],"genre_scores_gemma":[0.1341852,0.0008983448,0.8157163,0.0005332801,0.00014331537,0.0007720884,0.023484932,0.011618148,0.012648436],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99731123,0.00038317274,0.00019615544,0.0007493513,0.0012210102,0.00013909157],"domain_scores_gemma":[0.9911883,0.003385421,0.0005586025,0.0013114184,0.0031306532,0.00042558165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002464541,0.0019568335,0.0011231236,0.004722469,0.00090425415,0.0020016914,0.0022337586,0.0017374925,0.013032978],"category_scores_gemma":[0.021416482,0.0010208689,0.0012351837,0.0016764267,0.0003622164,0.0031076712,0.0013975598,0.0017817926,0.01272825],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053553766,0.0002543737,0.016011547,0.0008626267,0.0001025766,0.0004904683,0.00055939244,0.009540069,0.012612045,0.0059360955,0.2737517,0.67934346],"study_design_scores_gemma":[0.00025618842,0.00021751493,0.010698759,0.0003823729,0.00018035715,0.0010359165,0.0005550763,0.5956131,0.03881685,0.010598345,0.3414171,0.00022848777],"about_ca_topic_score_codex":0.01577902,"about_ca_topic_score_gemma":0.024618724,"teacher_disagreement_score":0.01577902,"about_ca_system_score_codex":0.0009360536,"about_ca_system_score_gemma":0.002758988,"threshold_uncertainty_score":0.043599665},"labels":[],"label_agreement":null},{"id":"W2137787810","doi":"10.1109/icse-companion.2009.5070977","title":"Detecting inefficient API usage","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Application programming interface; Java; Software; Code (set theory); Source code; Open source; Operating system; Software engineering; Database; World Wide Web; Programming language; Set (abstract data type)","score_opus":0.016424027092032657,"score_gpt":0.26478948053992785,"score_spread":0.24836545344789518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2137787810","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.74629474,0.0007197131,0.233145,0.00033300894,0.00009238034,0.00046940564,0.0018103234,0.011911456,0.005224084],"genre_scores_gemma":[0.8367663,0.00029645275,0.15501635,0.00018863524,0.000050896677,0.00042391638,0.0028565035,0.0016055572,0.002795384],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9851017,0.0031371855,0.0020745893,0.0026825082,0.006127313,0.00087672746],"domain_scores_gemma":[0.8986194,0.036363065,0.023697643,0.019167906,0.020611761,0.0015401577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005313004,0.0011143502,0.0012205831,0.007420898,0.0010151404,0.0018137167,0.0021090375,0.0013127669,0.0014847775],"category_scores_gemma":[0.067691356,0.00085585465,0.0009105929,0.0045377254,0.0006768468,0.0036707337,0.0021751358,0.0014289232,0.0011588582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007387289,0.00053054764,0.5372341,0.001132542,0.00047948517,0.001862876,0.0031873062,0.007969641,0.057614353,0.0052965055,0.010994046,0.37295988],"study_design_scores_gemma":[0.00013701002,0.0010188719,0.45432112,0.00042306748,0.0008811045,0.007233601,0.0022967705,0.33694375,0.14228222,0.012867065,0.041139293,0.0004560546],"about_ca_topic_score_codex":0.0022800711,"about_ca_topic_score_gemma":0.0029986454,"teacher_disagreement_score":0.007420898,"about_ca_system_score_codex":0.0007957636,"about_ca_system_score_gemma":0.0014355881,"threshold_uncertainty_score":0.028098226},"labels":[],"label_agreement":null},{"id":"W2138133615","doi":"10.1109/csmr.2012.51","title":"Using Topic Models to Support Software Maintenance","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software maintenance; Computer science; Software development; Topic model; Software engineering; Software; Data science; Information retrieval; Programming language","score_opus":0.09778445977607259,"score_gpt":0.3161165304156153,"score_spread":0.21833207063954274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138133615","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023783129,0.0018846514,0.9675822,0.0011866986,0.00015750974,0.00012255131,0.000688589,0.003243607,0.0013510679],"genre_scores_gemma":[0.5084494,0.0025332894,0.47916108,0.0004397386,0.0007323964,0.0008567482,0.0044469303,0.0006807857,0.0026997062],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99545693,0.0024696991,0.0002796811,0.0010130571,0.0005798521,0.00020079651],"domain_scores_gemma":[0.9547714,0.039105646,0.0019940068,0.0015604394,0.0020248422,0.0005437988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009683377,0.0015354911,0.0015692015,0.0044793175,0.0012272227,0.0031714782,0.0025295732,0.0021769395,0.0024379373],"category_scores_gemma":[0.04420688,0.0011090165,0.0021896742,0.0033759756,0.0009123064,0.0068394137,0.0024149558,0.0031798088,0.0013229991],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011050372,0.0005381742,0.020797247,0.0011629153,0.0010821574,0.00048060645,0.003237247,0.3946517,0.0038000308,0.06937798,0.019547597,0.48421922],"study_design_scores_gemma":[0.000062922,0.000053217325,0.0010031519,0.000040559145,0.00009592877,0.00007548603,0.00007554472,0.94820046,0.0004979789,0.046822965,0.0030388613,0.000032984604],"about_ca_topic_score_codex":0.009941731,"about_ca_topic_score_gemma":0.008759251,"teacher_disagreement_score":0.009941731,"about_ca_system_score_codex":0.0017011281,"about_ca_system_score_gemma":0.001720232,"threshold_uncertainty_score":0.05121118},"labels":[],"label_agreement":null},{"id":"W2138201395","doi":"10.1109/tse.2003.1214324","title":"An investigation of graph-based class integration test order strategies","year":2003,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":116,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Goddard Space Flight Center","keywords":"Computer science; Dependency (UML); Dependency graph; Inheritance (genetic algorithm); Class (philosophy); Theoretical computer science; Graph; Context (archaeology); Artificial intelligence","score_opus":0.01487302891763947,"score_gpt":0.2431641625133235,"score_spread":0.22829113359568404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138201395","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33523017,0.00090836495,0.6377694,0.00095376396,0.000036265603,0.0008063555,0.0001512122,0.0008616716,0.023282835],"genre_scores_gemma":[0.81799,0.00043343235,0.17971374,0.00010680978,0.000009754321,0.00017573885,0.00017209306,0.000095265954,0.001303199],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99653155,0.0016842732,0.00016407335,0.00026980243,0.0011468196,0.00020346597],"domain_scores_gemma":[0.97082055,0.023430567,0.0015584956,0.001302635,0.0025752753,0.0003124381],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004078192,0.00072744535,0.00051055785,0.0032901824,0.00049471005,0.0018573727,0.0013156252,0.0009544939,0.0031989156],"category_scores_gemma":[0.02720609,0.0003657544,0.0004998556,0.0018948094,0.0013624071,0.0033377542,0.0007291127,0.00077658525,0.0003514388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004125103,0.00085722894,0.017945087,0.00058026524,0.00010242437,0.0005240155,0.0020510873,0.24283512,0.013484426,0.2530182,0.0020188799,0.46617073],"study_design_scores_gemma":[0.00015173861,0.0008144247,0.0050106873,0.00012069332,0.0001247142,0.00045582379,0.0010220661,0.88074565,0.013110685,0.09126527,0.0071164225,0.0000619403],"about_ca_topic_score_codex":0.0054241186,"about_ca_topic_score_gemma":0.004907154,"teacher_disagreement_score":0.0054241186,"about_ca_system_score_codex":0.0022090296,"about_ca_system_score_gemma":0.0021081034,"threshold_uncertainty_score":0.021567822},"labels":[],"label_agreement":null},{"id":"W2138295189","doi":"10.1109/csmr.2008.4493326","title":"Visual Detection of Design Anomalies","year":2008,"lang":"en","type":"article","venue":"Proceedings of the ... European Conference on Software Maintenance and Reengineering/Proceedings of the European Conference on Software Maintenance and Reengineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Maintainability; Computer science; Anomaly detection; Flexibility (engineering); Visualization; Anomaly (physics); Task (project management); Data mining; Software; Artificial intelligence; Precision and recall; Machine learning; Software engineering; Engineering; Programming language; Systems engineering","score_opus":0.024622709699353186,"score_gpt":0.21746412173797972,"score_spread":0.19284141203862654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138295189","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13616534,0.0012668894,0.8084343,0.0007555162,0.00019846695,0.0004050707,0.00251741,0.03945933,0.0107976105],"genre_scores_gemma":[0.59664613,0.0006464816,0.39341947,0.00023269975,0.000095900396,0.00020687992,0.0024855759,0.0013238862,0.004943045],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99892646,0.00030241953,0.00006723963,0.00018172855,0.00044467254,0.000077405406],"domain_scores_gemma":[0.99080735,0.004312718,0.00093076524,0.0012100908,0.0024746475,0.0002645016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016545454,0.0011053865,0.00057006325,0.0040614414,0.0002577758,0.0014919058,0.00084526796,0.0009842943,0.0034775531],"category_scores_gemma":[0.01183355,0.00033099324,0.00040822712,0.0011551721,0.00026345375,0.00103882,0.0011403863,0.00078140665,0.0010411707],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011041579,0.00021914847,0.021096036,0.0013890936,0.00016433306,0.0012040449,0.0026826153,0.014233349,0.15546563,0.0067736865,0.029822543,0.7658454],"study_design_scores_gemma":[0.00028567857,0.0010643166,0.068776436,0.00074739696,0.0003584104,0.0054654474,0.00153025,0.5674264,0.2031732,0.028150702,0.12268723,0.00033462272],"about_ca_topic_score_codex":0.0013918212,"about_ca_topic_score_gemma":0.0012643544,"teacher_disagreement_score":0.0040614414,"about_ca_system_score_codex":0.000382443,"about_ca_system_score_gemma":0.00043285568,"threshold_uncertainty_score":0.011633515},"labels":[],"label_agreement":null},{"id":"W2138350998","doi":"10.1145/1518701.1518803","title":"A survey of software learnability","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":318,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Autodesk (Canada)","funders":"","keywords":"Learnability; Computer science; Usability; Protocol (science); Software; Think aloud protocol; Software engineering; Artificial intelligence; Human–computer interaction; Programming language; Medicine","score_opus":0.03215782087970254,"score_gpt":0.29247520620138173,"score_spread":0.2603173853216792,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138350998","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4330155,0.31612682,0.17367083,0.009809382,0.0005125336,0.0006814543,0.0027688616,0.0016171904,0.061797455],"genre_scores_gemma":[0.6966581,0.21767141,0.06846761,0.0012956199,0.0007027882,0.0007451548,0.0051663285,0.0004518179,0.008841193],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98612,0.0033105304,0.0023206917,0.0012093992,0.0065513127,0.00048817793],"domain_scores_gemma":[0.9243133,0.05236109,0.0048317173,0.0021072428,0.015242284,0.0011444748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009386824,0.00072821206,0.0010091398,0.016060615,0.0006937962,0.0025432599,0.0009715701,0.0010839604,0.0019188904],"category_scores_gemma":[0.0553945,0.00045530027,0.0007556652,0.013249845,0.0012209691,0.008887423,0.0015083947,0.0012530969,0.000562076],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013458625,0.00017610835,0.059222195,0.0037974254,0.000104366525,0.00015157287,0.0034720104,0.0018180595,0.0020274909,0.013358712,0.005022984,0.91071457],"study_design_scores_gemma":[0.00005044848,0.0025587953,0.408793,0.008211961,0.00025536967,0.0053170896,0.009164788,0.011553216,0.010390821,0.06900159,0.474357,0.00034588078],"about_ca_topic_score_codex":0.002259069,"about_ca_topic_score_gemma":0.0013802312,"teacher_disagreement_score":0.016060615,"about_ca_system_score_codex":0.0017041934,"about_ca_system_score_gemma":0.0015003341,"threshold_uncertainty_score":0.04964286},"labels":[],"label_agreement":null},{"id":"W2138427817","doi":"10.1145/2660190.2662114","title":"Does feature scattering follow power-law distributions?","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Fundação de Amparo à Pesquisa do Estado de Minas Gerais; Conselho Nacional de Desenvolvimento Científico e Tecnológico; Coordenação de Aperfeiçoamento de Pessoal de Nível Superior; Deutsche Forschungsgemeinschaft","keywords":"Scattering; Feature (linguistics); Computer science; Metric (unit); Code (set theory); Pareto distribution; Source code; Algorithm; Mathematics; Physics; Statistics; Optics; Engineering","score_opus":0.00804350452352562,"score_gpt":0.2411376878382516,"score_spread":0.23309418331472598,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138427817","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6231043,0.0009959674,0.36578497,0.0012203901,0.0000768246,0.00018257044,0.000592768,0.0016670655,0.006375231],"genre_scores_gemma":[0.98767626,0.000215979,0.010245759,0.00018494394,0.000045413144,0.000063857166,0.00039163162,0.00024066436,0.000935451],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.996014,0.0005993205,0.00022379705,0.0013042935,0.0013795412,0.00047898886],"domain_scores_gemma":[0.94940525,0.02744215,0.0091581885,0.007035874,0.0062099006,0.000748601],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00512431,0.00062570564,0.0009052589,0.002291883,0.0005941398,0.001829574,0.001623688,0.0011753562,0.0024546538],"category_scores_gemma":[0.051049788,0.0005720344,0.0008576514,0.0018759711,0.002527827,0.006000509,0.00093370606,0.0016017174,0.0010160567],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007420469,0.00042899756,0.53494585,0.0008143079,0.00053089845,0.002697174,0.005869702,0.14279258,0.02794108,0.103443205,0.011042459,0.16875175],"study_design_scores_gemma":[0.000059907507,0.00049902144,0.19215164,0.00016783027,0.0001411006,0.003196178,0.0021739833,0.58943427,0.013737148,0.18779218,0.010439465,0.00020723659],"about_ca_topic_score_codex":0.0029781586,"about_ca_topic_score_gemma":0.0023149874,"teacher_disagreement_score":0.00512431,"about_ca_system_score_codex":0.0012047305,"about_ca_system_score_gemma":0.00064950134,"threshold_uncertainty_score":0.027100325},"labels":[],"label_agreement":null},{"id":"W2138581775","doi":"","title":"An analysis of the design and definitions of halstead's metrics","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Set (abstract data type); Software; Data science; Software engineering; Programming language","score_opus":0.06426171022271657,"score_gpt":0.2900242928917718,"score_spread":0.22576258266905522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138581775","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030632457,0.0117769055,0.9080112,0.007481589,0.00058533694,0.0006662972,0.0003682495,0.0005936864,0.03988421],"genre_scores_gemma":[0.22634229,0.007387591,0.75919205,0.0011165226,0.00024650936,0.0011520396,0.00042213828,0.00030652704,0.0038343298],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96970797,0.011548734,0.0026401721,0.0017546208,0.013679742,0.00066880934],"domain_scores_gemma":[0.902707,0.057227124,0.00612173,0.0047206683,0.028423777,0.0007996342],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.022421606,0.0016486533,0.0008295319,0.012320906,0.0017277974,0.0044118515,0.001521865,0.0012197868,0.0012920691],"category_scores_gemma":[0.09217106,0.0008588614,0.0010876848,0.008304866,0.004986465,0.0090918485,0.0017609197,0.003264649,0.00040417368],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037109847,0.000050064173,0.0035791453,0.00040779024,0.000043492106,0.00009620725,0.0019938208,0.004549539,0.0016959434,0.7941958,0.0063194553,0.18703161],"study_design_scores_gemma":[0.00004248068,0.00062158704,0.015001729,0.0017520324,0.00013196001,0.0010600581,0.0024706651,0.034578566,0.013488531,0.6555143,0.27502397,0.00031409305],"about_ca_topic_score_codex":0.004926656,"about_ca_topic_score_gemma":0.0040611345,"teacher_disagreement_score":0.9775784,"about_ca_system_score_codex":0.006313533,"about_ca_system_score_gemma":0.0074303704,"threshold_uncertainty_score":0.118578255},"labels":[],"label_agreement":null},{"id":"W2138601625","doi":"10.5430/air.v4n2p45","title":"Automated selection of a software effort estimation model based on accuracy and uncertainty","year":2015,"lang":"en","type":"article","venue":"Artificial Intelligence Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Machine learning; Software; Selection (genetic algorithm); Estimation; Data mining; Process (computing); Model selection; Bayesian probability; Artificial intelligence; Software development; Engineering; Systems engineering","score_opus":0.14942728062778057,"score_gpt":0.4149915128714236,"score_spread":0.265564232243643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138601625","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.094997,0.00017675305,0.90250987,0.00025398948,0.00001288875,0.00012168988,0.00014802012,0.0009368266,0.0008429481],"genre_scores_gemma":[0.75870657,0.00016568274,0.23948014,0.0000772947,0.000027738233,0.00032826592,0.0007194151,0.00011220912,0.0003826015],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946243,0.0023960064,0.0004240798,0.0007439228,0.0015077682,0.0003038451],"domain_scores_gemma":[0.9800285,0.012809039,0.0020483206,0.0017435183,0.0031425348,0.00022803275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00791575,0.0011909334,0.0017805955,0.0040149167,0.0006856021,0.002438829,0.001531576,0.0011508354,0.0005811257],"category_scores_gemma":[0.031609032,0.00074811385,0.0012718436,0.0014269046,0.00058822543,0.002483883,0.0015344307,0.0014574677,0.0002905038],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026029657,0.00037730098,0.022866134,0.000192007,0.0002412353,0.00012448142,0.000275102,0.7824255,0.0074169426,0.0061415234,0.0014162483,0.17826314],"study_design_scores_gemma":[0.000013460179,0.000055915214,0.0022664138,0.000023239589,0.000028621733,0.000029904208,0.00003205259,0.9921665,0.0026694732,0.0023936338,0.00030067968,0.000020142861],"about_ca_topic_score_codex":0.0054870932,"about_ca_topic_score_gemma":0.00517097,"teacher_disagreement_score":0.00791575,"about_ca_system_score_codex":0.0013869943,"about_ca_system_score_gemma":0.002936547,"threshold_uncertainty_score":0.041862965},"labels":[],"label_agreement":null},{"id":"W2138702607","doi":"10.5555/381473.381485","title":"On the syllogistic structure of object-oriented programming","year":2001,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Syllogism; Theoretical computer science; Class hierarchy; Hierarchy; Class (philosophy); Programming language; Graph; Object-oriented programming; Artificial intelligence; Epistemology","score_opus":0.01604690833001546,"score_gpt":0.261478961741084,"score_spread":0.24543205341106852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138702607","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039492145,0.0027246124,0.83960485,0.010627768,0.00035887433,0.00007885731,0.0001834467,0.000803312,0.10612612],"genre_scores_gemma":[0.6682201,0.0028858765,0.3092413,0.0017788593,0.0008674993,0.000302564,0.0003218005,0.00039850414,0.01598351],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99769336,0.0010178887,0.0001239269,0.00025538314,0.0007095194,0.00019979026],"domain_scores_gemma":[0.99425215,0.0038885467,0.00032936642,0.00052490045,0.0006979583,0.00030702696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002651595,0.0003413629,0.00040118987,0.0017262766,0.002823378,0.0047349622,0.00076327665,0.001150226,0.004217242],"category_scores_gemma":[0.008622507,0.000579649,0.00060781406,0.0014081629,0.010799691,0.009234055,0.0027253325,0.0028883421,0.0007016968],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000004241931,0.000004717814,0.000099392906,0.000011743194,9.953643e-7,0.000019249828,0.0003097514,0.00043427036,0.000098341996,0.99395776,0.00056809804,0.004491429],"study_design_scores_gemma":[0.000003654334,0.0000041587596,0.00010467577,0.000016155785,0.0000020762618,0.000029395384,0.0000601739,0.0031926287,0.00012976705,0.986942,0.009508369,0.0000069739567],"about_ca_topic_score_codex":0.004824135,"about_ca_topic_score_gemma":0.0049217376,"teacher_disagreement_score":0.004824135,"about_ca_system_score_codex":0.0025621403,"about_ca_system_score_gemma":0.0024971743,"threshold_uncertainty_score":0.018589675},"labels":[],"label_agreement":null},{"id":"W2139101620","doi":"10.1109/apsec.2001.991506","title":"A framework for migrating procedural code to object-oriented platforms","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Legacy system; Computer science; Source code; Legacy code; Object-oriented programming; Business process reengineering; XML; Leverage (statistics); Code (set theory); Programming language; Software engineering; World Wide Web; Artificial intelligence; Engineering; Software","score_opus":0.024328610618189025,"score_gpt":0.3120893100294549,"score_spread":0.28776069941126586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139101620","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004896973,0.00008653277,0.9956311,0.0002462176,0.00005495209,0.00017470685,0.000027378335,0.002072484,0.0012168665],"genre_scores_gemma":[0.007411435,0.00026733123,0.9891928,0.00010874469,0.000046450372,0.00032647917,0.00015303976,0.0005378729,0.0019558894],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9960912,0.0010521606,0.00053911423,0.00047970418,0.0014504024,0.00038733703],"domain_scores_gemma":[0.9957684,0.0010354632,0.0004113387,0.0015142042,0.0007083085,0.0005623534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008504268,0.0015132359,0.0010616194,0.00307675,0.0031075678,0.007966429,0.007829852,0.004316243,0.0037662867],"category_scores_gemma":[0.009634262,0.002202937,0.004503187,0.0018735406,0.0064852796,0.005559842,0.0065322584,0.006257327,0.0030827245],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038198574,0.00012633673,0.00047326088,0.0003241216,0.000062576415,0.0007884722,0.0019067422,0.03313001,0.0051499223,0.87237775,0.0075188517,0.078103706],"study_design_scores_gemma":[0.000113249946,0.00016796116,0.0003440875,0.00066740386,0.00012490233,0.0012223509,0.0004913043,0.19756024,0.007946215,0.3857409,0.4053702,0.00025114842],"about_ca_topic_score_codex":0.011483008,"about_ca_topic_score_gemma":0.0108534545,"teacher_disagreement_score":0.011483008,"about_ca_system_score_codex":0.0021084736,"about_ca_system_score_gemma":0.0063790353,"threshold_uncertainty_score":0.0449754},"labels":[],"label_agreement":null},{"id":"W2139175514","doi":"10.1007/978-94-010-0421-3_26","title":"Reverse Engineering Interaction Plans for Legacy Interface Migration","year":2002,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Reverse engineering; Interface (matter); Legacy system; Computer science; Geology; Operating system","score_opus":0.03175177926017526,"score_gpt":0.26352455956108733,"score_spread":0.23177278030091208,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139175514","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011300373,0.00047027334,0.9390832,0.00041122723,0.0001209372,0.00017834827,0.00011812487,0.0034976387,0.044819888],"genre_scores_gemma":[0.1053789,0.0007397836,0.8350393,0.00020667791,0.000044370045,0.00021687169,0.0005698343,0.0015394508,0.056264795],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99907625,0.0002478058,0.00005686649,0.000106194,0.0004273641,0.00008561644],"domain_scores_gemma":[0.99816763,0.0007521153,0.000081553015,0.00062954926,0.00032893362,0.000040250536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013627562,0.0007532286,0.00033754358,0.0007468541,0.0009541988,0.0020345903,0.0013281798,0.0010745038,0.012313891],"category_scores_gemma":[0.0045542847,0.0007374342,0.00071644096,0.0006437215,0.00087326404,0.0029776872,0.0013609184,0.0021419504,0.00310522],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013613376,0.00019597673,0.0010110247,0.00025617532,0.000023038172,0.00042704257,0.0013183714,0.015905684,0.0101507995,0.2537512,0.023875132,0.6929494],"study_design_scores_gemma":[0.00007814317,0.00027936153,0.0011055355,0.0004584131,0.00017302846,0.0018125414,0.0010788774,0.30029193,0.055742208,0.35272723,0.28614807,0.00010456494],"about_ca_topic_score_codex":0.0021427853,"about_ca_topic_score_gemma":0.0036441789,"teacher_disagreement_score":0.012313891,"about_ca_system_score_codex":0.0005494632,"about_ca_system_score_gemma":0.0012445564,"threshold_uncertainty_score":0.04119408},"labels":[],"label_agreement":null},{"id":"W2139248042","doi":"10.1145/1082983.1083161","title":"SCQL","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Source code; Programming language; Software engineering; Software; Database; World Wide Web","score_opus":0.01701421333253953,"score_gpt":0.25024172086578145,"score_spread":0.23322750753324192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139248042","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009558643,0.00074248697,0.6233081,0.0030894764,0.00043074667,0.0014505475,0.07515009,0.23241664,0.053853247],"genre_scores_gemma":[0.21779849,0.0019102417,0.37783298,0.0057621542,0.00059444155,0.003061029,0.29119003,0.044165708,0.057684936],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99353504,0.0011764682,0.0009995179,0.0010656498,0.00268738,0.0005359695],"domain_scores_gemma":[0.9878907,0.0043608667,0.0006315463,0.0034390457,0.0032720184,0.00040574413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066364743,0.0013454951,0.0013039967,0.0029034382,0.0016955425,0.0065282183,0.004099374,0.002071689,0.039130144],"category_scores_gemma":[0.019732414,0.001068176,0.0020903193,0.0036281075,0.0014310675,0.009551646,0.0055961283,0.0025164017,0.017180799],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013197933,0.00027546927,0.0065373927,0.0020272303,0.00025879918,0.00068253005,0.0016917223,0.01143809,0.007891597,0.25683922,0.5224003,0.18863787],"study_design_scores_gemma":[0.0003097203,0.00013852499,0.0014220214,0.00029281306,0.00009080013,0.0005800464,0.00088507857,0.08553366,0.010986152,0.17074233,0.7288452,0.00017364953],"about_ca_topic_score_codex":0.015474018,"about_ca_topic_score_gemma":0.011188836,"teacher_disagreement_score":0.039130144,"about_ca_system_score_codex":0.0030315393,"about_ca_system_score_gemma":0.0036974892,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2139435518","doi":"10.1007/s11219-007-9015-6","title":"Introduction to the special issue on: “Software Quality Improvements and Estimations with Intelligence-based Methods”","year":2007,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software quality; Quality (philosophy); Software engineering; Software; Software development; Programming language","score_opus":0.0479755089915145,"score_gpt":0.39931576881835573,"score_spread":0.35134025982684125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139435518","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00089013804,0.044713646,0.04696248,0.034650303,0.83909845,0.00018881907,0.0006694189,0.0007321283,0.032094695],"genre_scores_gemma":[0.0064245537,0.040626433,0.015337167,0.024982845,0.79600865,0.00024510417,0.0018734016,0.0011264335,0.113375485],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99756646,0.0006156349,0.00034946427,0.00037898243,0.0009795127,0.000109845576],"domain_scores_gemma":[0.9890844,0.0051997183,0.0005342616,0.0006555425,0.0038924206,0.0006337173],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002417604,0.0021703814,0.002628638,0.0038975263,0.000985472,0.004180586,0.0020213316,0.003768035,0.05417098],"category_scores_gemma":[0.012641133,0.0007333988,0.0017041537,0.002607588,0.0010777372,0.004654613,0.0020147525,0.004654896,0.03288099],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039054874,0.000035768124,0.00013065581,0.00041273132,0.000028772827,0.00007272854,0.000025319414,0.00021244575,0.0004576491,0.0029817333,0.9407628,0.054840237],"study_design_scores_gemma":[0.00002244094,0.00009196706,0.000938574,0.00030733054,0.000045652425,0.00031130068,0.00004091949,0.0010744218,0.00047230843,0.0074244947,0.9892273,0.00004323163],"about_ca_topic_score_codex":0.00065058086,"about_ca_topic_score_gemma":0.001401218,"teacher_disagreement_score":0.05417098,"about_ca_system_score_codex":0.00083839195,"about_ca_system_score_gemma":0.0010417441,"threshold_uncertainty_score":0.18122},"labels":[],"label_agreement":null},{"id":"W2139440414","doi":"10.1109/apsec.2005.92","title":"Providing quality measurement for aspect-oriented software development","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Aspect-oriented programming; Computer science; Software engineering; Software development; Software development process; Software quality; Quality (philosophy); Process (computing); Software; Software construction; Systems engineering; Engineering; Programming language","score_opus":0.08114725363706193,"score_gpt":0.3178081296939096,"score_spread":0.23666087605684766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139440414","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06176545,0.00066575,0.93163943,0.00038311468,0.00011239579,0.00027729492,0.00023964256,0.0017188865,0.003197981],"genre_scores_gemma":[0.55256784,0.00046133486,0.4450338,0.000098538265,0.000083765226,0.0005953623,0.00064021914,0.0002052126,0.0003139748],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96841127,0.010909519,0.0038414418,0.0019428578,0.0141790435,0.00071598316],"domain_scores_gemma":[0.89270884,0.0416707,0.014367543,0.017585024,0.032070402,0.0015975892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015417297,0.0012422742,0.0013311957,0.007959532,0.00098721,0.003886654,0.0015685457,0.0017130384,0.0006437077],"category_scores_gemma":[0.10570168,0.00068679824,0.00074956496,0.0058721183,0.0012129176,0.0051259436,0.0022199356,0.0017792501,0.00038618452],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067980855,0.00088483986,0.089822836,0.0016005107,0.0004276487,0.00021115584,0.0025574225,0.0430516,0.04986633,0.05315081,0.0039314106,0.75381553],"study_design_scores_gemma":[0.00025141586,0.0021952516,0.08123867,0.0010395269,0.00043954127,0.0010099924,0.0018882023,0.6544917,0.14626437,0.08067833,0.02988863,0.0006143869],"about_ca_topic_score_codex":0.0012945912,"about_ca_topic_score_gemma":0.0009785406,"teacher_disagreement_score":0.015417297,"about_ca_system_score_codex":0.0013241823,"about_ca_system_score_gemma":0.0013210179,"threshold_uncertainty_score":0.0815354},"labels":[],"label_agreement":null},{"id":"W2139694402","doi":"10.1109/vissof.2002.1019798","title":"The CONCEPT project - applying source code analysis to reduce information complexity of static and dynamic visualization techniques","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Source code; Program slicing; Program comprehension; Visualization; Static program analysis; KPI-driven code analysis; Software visualization; Programming language; Slicing; Static analysis; Unified Modeling Language; Software engineering; Data mining; Software; Software system; Software development; Software construction; World Wide Web","score_opus":0.027367278237062014,"score_gpt":0.33555846318677685,"score_spread":0.3081911849497148,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139694402","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013389194,0.00017819744,0.9796053,0.0006315365,0.0000447398,0.00015866691,0.00004991923,0.0031546648,0.0027877863],"genre_scores_gemma":[0.09703441,0.00030211563,0.8994577,0.00014995199,0.000049661205,0.0003034885,0.00019356488,0.00064726366,0.0018618745],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9970138,0.0012733313,0.00012203556,0.00034579774,0.0011052444,0.00013976509],"domain_scores_gemma":[0.9885576,0.0062381914,0.00055388454,0.0019162077,0.0024571484,0.0002769687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046327272,0.0014557941,0.00073004444,0.0028702875,0.0008077938,0.003398831,0.0020178738,0.0013343742,0.0055159926],"category_scores_gemma":[0.017527694,0.00071114727,0.001048332,0.002151459,0.0019334458,0.00912813,0.002869926,0.0017176078,0.0010129908],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041136707,0.00044016735,0.0032182653,0.0010314686,0.00014367238,0.00047480915,0.0046521584,0.01252981,0.070445865,0.16257322,0.01038429,0.7336949],"study_design_scores_gemma":[0.0003318047,0.0013764128,0.004086913,0.0006031159,0.00028721878,0.0029086769,0.0019498772,0.40070176,0.18347217,0.25234672,0.15162684,0.00030859542],"about_ca_topic_score_codex":0.0012641113,"about_ca_topic_score_gemma":0.00068086677,"teacher_disagreement_score":0.0055159926,"about_ca_system_score_codex":0.00058842334,"about_ca_system_score_gemma":0.0018120757,"threshold_uncertainty_score":0.02450049},"labels":[],"label_agreement":null},{"id":"W2139776687","doi":"10.1109/ase.2002.1115033","title":"Predicting software stability using case-based reasoning","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software metric; Software sizing; Software; Software construction; Software development; Artificial intelligence; Search-based software engineering; Software maintenance; Machine learning; Data mining; Software engineering; Programming language","score_opus":0.056757058708049896,"score_gpt":0.2717961779794438,"score_spread":0.21503911927139388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139776687","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6974392,0.00059611874,0.2930022,0.00085499906,0.000051068815,0.000542105,0.0017777972,0.0013789923,0.004357604],"genre_scores_gemma":[0.8886532,0.00018794458,0.10870124,0.00004507071,0.00003772851,0.00014098379,0.0018809977,0.000021192638,0.00033168306],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99553275,0.001505965,0.00055971695,0.0006733146,0.0015602867,0.00016803016],"domain_scores_gemma":[0.96176165,0.031353563,0.002521529,0.0016851783,0.0021878176,0.00049040484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004192912,0.00093994616,0.0007115765,0.009675249,0.0007244422,0.002052079,0.0015686398,0.0021296607,0.0016294189],"category_scores_gemma":[0.038744666,0.00037508926,0.0008058083,0.0042883074,0.00087252556,0.0029330368,0.00081593025,0.0007296504,0.00048515486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063951174,0.0013827413,0.19569752,0.00027820043,0.00040698913,0.00117118,0.00046959118,0.3905565,0.0034086334,0.0072760074,0.0041935747,0.39451957],"study_design_scores_gemma":[0.000041883384,0.00007667933,0.007756045,0.000020487416,0.000055941066,0.00018809571,0.00007986387,0.98040056,0.0016226938,0.009207963,0.0005224592,0.000027273642],"about_ca_topic_score_codex":0.005324136,"about_ca_topic_score_gemma":0.0050169183,"teacher_disagreement_score":0.009675249,"about_ca_system_score_codex":0.0011055009,"about_ca_system_score_gemma":0.00082679995,"threshold_uncertainty_score":0.022174537},"labels":[],"label_agreement":null},{"id":"W2140037122","doi":"10.1109/kbse.1992.252905","title":"CAESAR: a system for case based software reuse","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Reuse; Software; Software engineering; Representation (politics); Software system; Software development; Artificial intelligence; Programming language; Engineering","score_opus":0.02417503503063203,"score_gpt":0.26398116056142246,"score_spread":0.23980612553079042,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140037122","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015297335,0.0001834869,0.9020409,0.0002443026,0.00011366295,0.00051724236,0.00066696294,0.08714907,0.0075546433],"genre_scores_gemma":[0.023545276,0.00046824175,0.9552739,0.00028773627,0.000086467546,0.0007824005,0.004117244,0.0054163057,0.010022351],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99513525,0.0011979039,0.00080913666,0.0008144122,0.0018331766,0.00021010128],"domain_scores_gemma":[0.9921261,0.0035484696,0.00037733658,0.0027984183,0.0008769649,0.0002727026],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052768895,0.0015283601,0.0013708057,0.0066123763,0.0011923902,0.005515569,0.0046310388,0.00266943,0.027407335],"category_scores_gemma":[0.01906091,0.0024379953,0.0026272086,0.0038089857,0.0018730162,0.006935675,0.006897244,0.0030838458,0.01735432],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051402027,0.0004686984,0.0018263775,0.0009818794,0.00031052635,0.0014339202,0.0016688526,0.020927418,0.01062,0.20954959,0.12021817,0.63148063],"study_design_scores_gemma":[0.0004532445,0.00017491788,0.0008667007,0.00043582235,0.00028611734,0.0020150125,0.00026670622,0.29608166,0.024045197,0.16499904,0.5100767,0.000298883],"about_ca_topic_score_codex":0.002657147,"about_ca_topic_score_gemma":0.0025379078,"teacher_disagreement_score":0.027407335,"about_ca_system_score_codex":0.001098008,"about_ca_system_score_gemma":0.0018678833,"threshold_uncertainty_score":0.091686726},"labels":[],"label_agreement":null},{"id":"W2140113450","doi":"10.3844/jcssp.2008.571.577","title":"An Empirical Validation of Object-Oriented Design Metrics for Fault Prediction","year":2008,"lang":"en","type":"article","venue":"Journal of Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Empirical research; Fault (geology); Data mining; Artificial intelligence; Statistics","score_opus":0.06229142806206331,"score_gpt":0.33621341128855825,"score_spread":0.27392198322649497,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140113450","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99217975,0.00010177776,0.0063141095,0.000069853435,0.000011396672,0.00007781752,0.0003212244,0.00004623075,0.00087786466],"genre_scores_gemma":[0.99518096,0.000025592599,0.0040898696,0.000012041119,0.0000059008494,0.00007048917,0.0005179772,0.0000061487003,0.00009096177],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9888691,0.0058796653,0.0007620955,0.0006291408,0.0036935161,0.0001665334],"domain_scores_gemma":[0.85751975,0.08438575,0.015928095,0.008980663,0.03206874,0.0011169981],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012333277,0.00051380834,0.00029570778,0.0031083243,0.00036652546,0.0007098517,0.0008959097,0.0007336742,0.0007420251],"category_scores_gemma":[0.0665335,0.00012892878,0.00036080962,0.0021041597,0.0005718635,0.001136597,0.00063154195,0.000514302,0.00029736423],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004891859,0.0022419707,0.8729711,0.0002449571,0.00023104741,0.00011082286,0.00074471213,0.019262895,0.0035921135,0.001064482,0.001560322,0.09748636],"study_design_scores_gemma":[0.00013124476,0.0037756606,0.72105545,0.0002028148,0.00015484658,0.0002779268,0.0010461344,0.2557354,0.013112435,0.0015085513,0.0029450522,0.00005434647],"about_ca_topic_score_codex":0.002265837,"about_ca_topic_score_gemma":0.002017604,"teacher_disagreement_score":0.012333277,"about_ca_system_score_codex":0.0009894421,"about_ca_system_score_gemma":0.0007705843,"threshold_uncertainty_score":0.06522536},"labels":[],"label_agreement":null},{"id":"W2140183398","doi":"10.1109/icsm.2011.6080777","title":"Generating natural language summaries for crosscutting source code concerns","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":81,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Programming language; Task (project management); Source code; Code (set theory); Code review; KPI-driven code analysis; Natural language; Software engineering; Software; Natural (archaeology); Software development; Static program analysis; Artificial intelligence; Systems engineering; Engineering","score_opus":0.04171048364221037,"score_gpt":0.3016687192181296,"score_spread":0.25995823557591924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140183398","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19148642,0.0005885363,0.7504991,0.0011805469,0.00020200823,0.0011683325,0.0076146573,0.044524536,0.0027357673],"genre_scores_gemma":[0.24778518,0.0002876629,0.7350759,0.0002150265,0.00006138772,0.00050341117,0.0123614315,0.0020779094,0.0016321337],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974025,0.0011563379,0.00026298506,0.00039033644,0.0007068666,0.00008096946],"domain_scores_gemma":[0.9692516,0.023315731,0.0026011702,0.0019298287,0.0026117952,0.00028990806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026044287,0.0015136718,0.0005677831,0.0016691596,0.0005627051,0.0011849046,0.0015309608,0.0011196886,0.002918012],"category_scores_gemma":[0.023221398,0.00064120314,0.0008067299,0.0008831523,0.00055198936,0.0020295177,0.0010845544,0.0013156572,0.0009941444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016939432,0.001204862,0.013013925,0.0050309706,0.00022878931,0.004362629,0.013350471,0.062556915,0.16500095,0.025075214,0.0635105,0.6449707],"study_design_scores_gemma":[0.0006820375,0.0013331564,0.0071525984,0.0003819759,0.00042143962,0.0027029922,0.003158438,0.6648974,0.16899687,0.032108407,0.11786115,0.00030354058],"about_ca_topic_score_codex":0.0027224065,"about_ca_topic_score_gemma":0.0059926216,"teacher_disagreement_score":0.002918012,"about_ca_system_score_codex":0.0007633276,"about_ca_system_score_gemma":0.0016811731,"threshold_uncertainty_score":0.01377368},"labels":[],"label_agreement":null},{"id":"W2140466350","doi":"10.1109/msr.2013.6624002","title":"Gerrit software code review data from Android","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software engineering; Android (operating system); Software development; Software development process; Software; JSON; Database; Programming language; Operating system","score_opus":0.062055572819611216,"score_gpt":0.30724426560621826,"score_spread":0.24518869278660704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140466350","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4384079,0.0048701367,0.020513851,0.0035513914,0.00053153746,0.0050395313,0.42103332,0.006998372,0.09905395],"genre_scores_gemma":[0.3804822,0.002968388,0.06357793,0.0010399253,0.00040182838,0.008991979,0.4978915,0.002469394,0.04217684],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9923745,0.00085916626,0.0014846987,0.00085880666,0.0041032336,0.00031963649],"domain_scores_gemma":[0.9318674,0.019374812,0.009200165,0.0083238175,0.029748017,0.0014857769],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0031418707,0.00043984147,0.0004118263,0.012741745,0.0009647351,0.0015508041,0.0007138988,0.0006726611,0.004940753],"category_scores_gemma":[0.043004997,0.00034492882,0.00041812554,0.007099349,0.00066390896,0.0010081591,0.0020267197,0.0009257015,0.0049486244],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008738783,0.00036440624,0.18197034,0.0061098915,0.00013520566,0.0023306443,0.013185294,0.0014987934,0.012920206,0.009733165,0.37751782,0.3933603],"study_design_scores_gemma":[0.00011185183,0.00019554692,0.2806713,0.0009268904,0.00007483033,0.0011375418,0.0027369724,0.0019324524,0.009927965,0.0017962741,0.7003664,0.000122021345],"about_ca_topic_score_codex":0.012682164,"about_ca_topic_score_gemma":0.01949919,"teacher_disagreement_score":0.9968581,"about_ca_system_score_codex":0.0012913635,"about_ca_system_score_gemma":0.0035426535,"threshold_uncertainty_score":0.025216699},"labels":[],"label_agreement":null},{"id":"W2140598376","doi":"10.1109/icsm.1995.526535","title":"Design maintenance: unexpected architectural interactions (experience report)","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; IBM (Canada)","funders":"","keywords":"Categorization; Computer science; Code (set theory); Software engineering; Systems design; Systems engineering; Programming language; Human–computer interaction; Engineering; Artificial intelligence; Set (abstract data type)","score_opus":0.0445186824504735,"score_gpt":0.2870970631887051,"score_spread":0.2425783807382316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140598376","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8075178,0.003174501,0.12575173,0.01013861,0.0006128539,0.000428112,0.0006991103,0.0074501354,0.04422713],"genre_scores_gemma":[0.9108607,0.0018486329,0.054852553,0.00244408,0.00023304822,0.0001887048,0.001404433,0.0015969675,0.026570873],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9914125,0.0031823143,0.0004390991,0.0008786007,0.0034773848,0.00061013794],"domain_scores_gemma":[0.9647948,0.017116902,0.0033308088,0.0069636973,0.00573234,0.0020614571],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009642468,0.0009414815,0.0003849888,0.0017650343,0.0022951455,0.0025963464,0.002401474,0.0026822889,0.0040534735],"category_scores_gemma":[0.046823457,0.00086497323,0.0007295183,0.0013080282,0.0019299898,0.0030082217,0.0025565287,0.0028026656,0.001575854],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063691376,0.0029514502,0.049452268,0.0011016186,0.00021759092,0.012697118,0.09968604,0.0054323236,0.01919585,0.005817204,0.09064894,0.7121627],"study_design_scores_gemma":[0.00049379235,0.0057577924,0.07705544,0.0010758265,0.00072374445,0.045751993,0.037484623,0.025311213,0.05869979,0.009925593,0.7371592,0.0005609313],"about_ca_topic_score_codex":0.003670979,"about_ca_topic_score_gemma":0.009486152,"teacher_disagreement_score":0.009642468,"about_ca_system_score_codex":0.0013586262,"about_ca_system_score_gemma":0.0015429424,"threshold_uncertainty_score":0.050994873},"labels":[],"label_agreement":null},{"id":"W2140629797","doi":"10.1007/s00766-012-0159-y","title":"Guidelines for using UML association classes and their effect on domain understanding in requirements engineering","year":2012,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Unified Modeling Language; Computer science; Software engineering; Applications of UML; UML tool; Domain (mathematical analysis); Association (psychology); Systems engineering; Programming language; Engineering; Psychology; Mathematics; Software","score_opus":0.14466363731676277,"score_gpt":0.35794081310961684,"score_spread":0.21327717579285407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140629797","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03273638,0.0013864823,0.90543157,0.015356666,0.00045241055,0.0020369596,0.0005157239,0.008954266,0.033129625],"genre_scores_gemma":[0.045785,0.0005686163,0.9485602,0.0009260076,0.00006425382,0.00071421126,0.0002915402,0.0008409665,0.0022492204],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9443267,0.032842506,0.008682106,0.0014647505,0.01179255,0.0008913197],"domain_scores_gemma":[0.627517,0.26761803,0.016550548,0.021714695,0.063626416,0.0029733048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.042499844,0.001042246,0.0007204758,0.005687169,0.0023603388,0.0056467634,0.0026765142,0.0050076037,0.004797063],"category_scores_gemma":[0.21339989,0.002260242,0.00093872356,0.0039795022,0.0022749817,0.0069467453,0.0028830967,0.006156886,0.0026697794],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000569207,0.001468495,0.0121473055,0.002386899,0.00009483925,0.0009257938,0.016894668,0.0056800297,0.022758227,0.15565987,0.056989413,0.72442526],"study_design_scores_gemma":[0.0008917775,0.0009588842,0.023275824,0.013915565,0.0007489344,0.003418192,0.008151316,0.07744658,0.06762795,0.27195033,0.5307807,0.0008339157],"about_ca_topic_score_codex":0.0060820053,"about_ca_topic_score_gemma":0.02079687,"teacher_disagreement_score":0.042499844,"about_ca_system_score_codex":0.0015486925,"about_ca_system_score_gemma":0.004984036,"threshold_uncertainty_score":0.22476333},"labels":[],"label_agreement":null},{"id":"W2140697933","doi":"10.1109/wcre.2006.28","title":"Extracting Facts from Perl Code","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Perl; Computer science; Scripting language; Programming language; Extractor; Reverse engineering; Interpreter; Python (programming language); Source code; Software engineering","score_opus":0.021802170045696952,"score_gpt":0.26513402569679684,"score_spread":0.24333185565109988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140697933","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054671004,0.00031869893,0.8658788,0.00074488134,0.00013488553,0.0008313889,0.011959306,0.04574965,0.01971134],"genre_scores_gemma":[0.16214898,0.00056551985,0.79627866,0.00017699061,0.00007521892,0.00042987784,0.024234496,0.0055193226,0.0105708],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986533,0.0001794566,0.00011747924,0.000255872,0.000706236,0.00008761394],"domain_scores_gemma":[0.9944595,0.0022700143,0.00044875054,0.0012541333,0.0014941714,0.000073456875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014214206,0.000942862,0.00046313036,0.0036599822,0.00054678175,0.0015981786,0.0008881688,0.0005418306,0.0074132737],"category_scores_gemma":[0.007543416,0.0005818928,0.00083401834,0.0018453618,0.0007055643,0.002208044,0.0011272121,0.0009025055,0.0031615125],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040890736,0.00018758765,0.008860338,0.0022544602,0.00009384979,0.0041360618,0.003465859,0.013682147,0.068488,0.05269937,0.03482445,0.8108991],"study_design_scores_gemma":[0.00009692664,0.00027623496,0.010391283,0.0006061446,0.00019549456,0.0038044665,0.0011303513,0.20912652,0.3024362,0.050490167,0.42128786,0.00015844636],"about_ca_topic_score_codex":0.0011372339,"about_ca_topic_score_gemma":0.0018094538,"teacher_disagreement_score":0.0074132737,"about_ca_system_score_codex":0.0007934119,"about_ca_system_score_gemma":0.0013612023,"threshold_uncertainty_score":0.024799824},"labels":[],"label_agreement":null},{"id":"W2140921887","doi":"10.1109/ccece.2003.1226140","title":"The utility of graph theoretic software metrics: a case study","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Institute for Biodiagnostics","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software sizing; Software; Software metric; Call graph; Theoretical computer science; Software construction; Software development; Inheritance (genetic algorithm); Class (philosophy); Graph; Software visualization; Data mining; Software engineering; Object-oriented programming; Programming language; Artificial intelligence","score_opus":0.023573536483452156,"score_gpt":0.288105597269525,"score_spread":0.26453206078607283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140921887","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8915313,0.001344652,0.08236028,0.004264119,0.000048624028,0.00043186627,0.00042393347,0.000323797,0.019271385],"genre_scores_gemma":[0.9282471,0.0007430938,0.06861665,0.00017938095,0.000025693385,0.00016303301,0.00017259373,0.00012577661,0.0017267647],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9908901,0.00695253,0.00024777645,0.0002863125,0.001378302,0.00024494398],"domain_scores_gemma":[0.9407449,0.05094388,0.0016884698,0.0025342607,0.0032320619,0.0008564491],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00963413,0.0008220502,0.0006059697,0.0034003041,0.0021278993,0.0022340042,0.0018468009,0.002967574,0.0012963383],"category_scores_gemma":[0.03234647,0.0002953218,0.0006585647,0.0055221035,0.003267072,0.0032651005,0.0019644175,0.0017776925,0.00020947185],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010430081,0.0039590406,0.101330854,0.0015609695,0.00024524485,0.028306823,0.041812226,0.11592082,0.009424111,0.31821415,0.01980229,0.35838044],"study_design_scores_gemma":[0.000578343,0.0036274109,0.062497772,0.00086125865,0.0003142694,0.017314754,0.04769014,0.5019792,0.019585745,0.24638653,0.09881646,0.0003481675],"about_ca_topic_score_codex":0.008869536,"about_ca_topic_score_gemma":0.017246826,"teacher_disagreement_score":0.00963413,"about_ca_system_score_codex":0.0029360224,"about_ca_system_score_gemma":0.0011200758,"threshold_uncertainty_score":0.050950706},"labels":[],"label_agreement":null},{"id":"W2140972380","doi":"10.1145/1985793.1986025","title":"Measuring subversions","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Java; Software; Set (abstract data type); Software engineering; Secure coding; Software system; Term (time); Security bug; Vulnerability (computing); Backporting; World Wide Web; Software security assurance; Computer security; Software construction; Operating system; Information security; Programming language","score_opus":0.11213437236226843,"score_gpt":0.24224478949758635,"score_spread":0.1301104171353179,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140972380","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.932415,0.0011263144,0.036109664,0.00014339841,0.00012340643,0.00027478553,0.0037156558,0.0024255912,0.023666278],"genre_scores_gemma":[0.9645314,0.0004223757,0.026675792,0.00004912663,0.000044931076,0.00014700585,0.0039529167,0.00048155742,0.0036949427],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9924045,0.0013653963,0.0012801965,0.001227865,0.0033958391,0.000326242],"domain_scores_gemma":[0.9221691,0.032387868,0.011703146,0.014774677,0.01636581,0.0025994764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042911386,0.0007856843,0.00046572063,0.007843671,0.0006020244,0.0024549062,0.0010006455,0.0006199471,0.0028369373],"category_scores_gemma":[0.054769676,0.00044193148,0.000676919,0.004765463,0.0006461414,0.0028991164,0.0016476365,0.001106234,0.001792162],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038332006,0.00026695462,0.6367371,0.00077168166,0.00025637276,0.000387336,0.0076737385,0.004050946,0.014665245,0.0047634947,0.004055969,0.32598782],"study_design_scores_gemma":[0.000035233137,0.0006536893,0.912316,0.00020625436,0.00027785866,0.002363664,0.0020832184,0.0137111135,0.025191398,0.0063630478,0.036655392,0.00014307522],"about_ca_topic_score_codex":0.0018510188,"about_ca_topic_score_gemma":0.002019137,"teacher_disagreement_score":0.007843671,"about_ca_system_score_codex":0.0008395139,"about_ca_system_score_gemma":0.0006647168,"threshold_uncertainty_score":0.022693932},"labels":[],"label_agreement":null},{"id":"W2141048470","doi":"10.1109/icsm.2002.1167814","title":"Migration to object oriented platforms: a state transformation approach","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Legacy system; Business process reengineering; Software engineering; Model transformation; Sequence diagram; Workbench; Context (archaeology); Programming language; Object-oriented programming; Legacy code; Process (computing); Software quality; Systems engineering; Software development; Unified Modeling Language; Software; Engineering; Data mining; Artificial intelligence; Visualization","score_opus":0.015405384559831528,"score_gpt":0.24472580809215677,"score_spread":0.22932042353232523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141048470","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00977919,0.00005473755,0.98668563,0.00017336747,0.000018295099,0.000050315626,0.000032253574,0.0007015834,0.0025046319],"genre_scores_gemma":[0.47092906,0.00045459755,0.5201704,0.00014822032,0.000041647247,0.00038412533,0.00022976557,0.00023196562,0.0074102497],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99948525,0.00021397819,0.000029052462,0.00009010103,0.00013067233,0.00005088674],"domain_scores_gemma":[0.9990771,0.00049507275,0.00010681957,0.00018446613,0.00010359896,0.00003292401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00077711715,0.00037379045,0.00041003022,0.00063591625,0.0005593824,0.0012905938,0.000996666,0.0008191203,0.002723305],"category_scores_gemma":[0.002738559,0.00034182158,0.00084354775,0.0005856789,0.0010503723,0.0015285369,0.0012273168,0.0009904428,0.0005673814],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012807873,0.0001704273,0.0010694956,0.000108529544,0.000049220915,0.00022588478,0.0004475431,0.6257616,0.011112407,0.23115057,0.0013807962,0.12839545],"study_design_scores_gemma":[0.00001348248,0.000043098124,0.00011629017,0.000009948376,0.00001391269,0.000030961404,0.000031906293,0.94432604,0.0042835227,0.048190035,0.002927392,0.000013334864],"about_ca_topic_score_codex":0.0034131128,"about_ca_topic_score_gemma":0.0029649495,"teacher_disagreement_score":0.0034131128,"about_ca_system_score_codex":0.0008741282,"about_ca_system_score_gemma":0.0010875987,"threshold_uncertainty_score":0.009110391},"labels":[],"label_agreement":null},{"id":"W2141156769","doi":"10.1109/icsm.2002.1167758","title":"Behavioural concern modelling for software change tasks","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"University of British Columbia","keywords":"Computer science; Software engineering; Code (set theory); Software; Human–computer interaction; Finite-state machine; Programming language","score_opus":0.1649159748589482,"score_gpt":0.3132508083726179,"score_spread":0.14833483351366966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141156769","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003626314,0.0000884809,0.99139386,0.00034994472,0.000022393395,0.00013693358,0.00009609585,0.00048416693,0.0038018033],"genre_scores_gemma":[0.1682707,0.0003258985,0.82370555,0.00027944526,0.00006628412,0.0010791827,0.00067693536,0.00036276339,0.0052332045],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9946732,0.002610201,0.0004960405,0.0006498404,0.0012298858,0.00034096136],"domain_scores_gemma":[0.9905691,0.005662314,0.00087010727,0.0015404093,0.0010801698,0.0002779752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046523865,0.0015932727,0.0005762821,0.0022531352,0.0010511857,0.0027846447,0.002320272,0.0028231726,0.004180552],"category_scores_gemma":[0.01887354,0.0011351386,0.0030807557,0.001388715,0.0030331337,0.0053731645,0.0023782165,0.0032534515,0.001358643],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009410512,0.00012919417,0.0024098733,0.00041419148,0.000070109,0.00036375583,0.0030509199,0.11896073,0.0038519565,0.82174104,0.0028122854,0.046101842],"study_design_scores_gemma":[0.000042536045,0.00006604424,0.00050253083,0.0001355008,0.000062829204,0.00022472296,0.00031658617,0.43847075,0.002674125,0.5255929,0.031853914,0.000057549296],"about_ca_topic_score_codex":0.007827478,"about_ca_topic_score_gemma":0.008297866,"teacher_disagreement_score":0.007827478,"about_ca_system_score_codex":0.0024277882,"about_ca_system_score_gemma":0.0027207434,"threshold_uncertainty_score":0.02460444},"labels":[],"label_agreement":null},{"id":"W2141506510","doi":"10.1002/smr.299","title":"Improving design quality using meta‐pattern transformations: a metric‐based approach","year":2004,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Metric (unit); Business process reengineering; Task (project management); Quality (philosophy); Process (computing); Software engineering; Object-oriented design; Metamodeling; Data mining; Software; Systems engineering; Programming language; Engineering","score_opus":0.17194493028079816,"score_gpt":0.38775112576598647,"score_spread":0.2158061954851883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141506510","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046840888,0.0004928652,0.9464494,0.00068088085,0.000028861492,0.00022502457,0.00012096799,0.0025447938,0.0026163622],"genre_scores_gemma":[0.3904776,0.00030797467,0.607563,0.0000884833,0.00002520997,0.0002437461,0.0002885106,0.00038606804,0.00061943586],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9904998,0.0036847845,0.00072883005,0.00077060517,0.004108985,0.00020703247],"domain_scores_gemma":[0.9717575,0.009906155,0.005408304,0.006134574,0.0062896954,0.00050377985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071350387,0.0012680612,0.0011215194,0.0054306975,0.0004949001,0.00282613,0.0021103122,0.00095681194,0.00085472583],"category_scores_gemma":[0.029578121,0.00062456646,0.0008972021,0.0033269043,0.0015236144,0.003347207,0.0015449749,0.0013918298,0.00022776776],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031347352,0.00054293586,0.019999208,0.00073557644,0.0003431272,0.0002490884,0.0008219444,0.12323554,0.03620781,0.06352666,0.002896731,0.75112784],"study_design_scores_gemma":[0.000099275836,0.0006594318,0.010610827,0.00019322305,0.00023974455,0.0005887777,0.00029723617,0.8633794,0.043525025,0.07051177,0.009774994,0.000120275414],"about_ca_topic_score_codex":0.0015643798,"about_ca_topic_score_gemma":0.0016082588,"teacher_disagreement_score":0.0071350387,"about_ca_system_score_codex":0.0016781219,"about_ca_system_score_gemma":0.0018892159,"threshold_uncertainty_score":0.03773409},"labels":[],"label_agreement":null},{"id":"W2141519322","doi":"10.5555/2666990.2666993","title":"Usability challenges in exception handling","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Assertion; Exception handling; Computer science; Usability; Programming language; Matching (statistics); Block (permutation group theory); Human–computer interaction; Mathematics","score_opus":0.09250696725919971,"score_gpt":0.30998267618308695,"score_spread":0.21747570892388723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141519322","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16887023,0.007965681,0.71456707,0.039986618,0.001554208,0.0015478315,0.00027564212,0.017096547,0.048136175],"genre_scores_gemma":[0.53058696,0.0030875248,0.43757072,0.0060080728,0.00086680334,0.00075090484,0.00038163437,0.0063823685,0.014365032],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.8546676,0.08629761,0.011245122,0.007351369,0.037167512,0.0032708074],"domain_scores_gemma":[0.5442146,0.31233153,0.012777992,0.057765223,0.06918514,0.0037255574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10867852,0.0014469415,0.0014217177,0.0036426852,0.0033922214,0.017914945,0.006771242,0.0040820907,0.005594454],"category_scores_gemma":[0.27455375,0.0014096717,0.001653318,0.002485113,0.0064251553,0.017007848,0.00720362,0.005084015,0.0026400439],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010325505,0.0005931404,0.018863019,0.00531997,0.00021953104,0.0026529701,0.06009917,0.003758725,0.021424739,0.0836159,0.028085196,0.7743351],"study_design_scores_gemma":[0.0004720915,0.003678015,0.015026778,0.01042371,0.000828672,0.015493914,0.06483206,0.05244597,0.033540953,0.30530727,0.49668103,0.001269559],"about_ca_topic_score_codex":0.0050204378,"about_ca_topic_score_gemma":0.004451597,"teacher_disagreement_score":0.10867852,"about_ca_system_score_codex":0.0032042435,"about_ca_system_score_gemma":0.00580226,"threshold_uncertainty_score":0.5747538},"labels":[],"label_agreement":null},{"id":"W2142121021","doi":"10.1109/icse-companion.2009.5071003","title":"How do system architectures affect software requirements?","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software requirements specification; System requirements; Software requirements; Requirements analysis; Software architecture; Software engineering; Reference architecture; Architecture; Software; Systems architecture; Software system; Systems engineering; Software construction; Engineering; Operating system","score_opus":0.01851258310073587,"score_gpt":0.26465502946416675,"score_spread":0.24614244636343088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142121021","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9751881,0.00013435332,0.00458877,0.00071813364,0.000010714026,0.00005922264,0.00004846794,0.000044550165,0.019207653],"genre_scores_gemma":[0.9986066,0.000046308338,0.0009673778,0.000063419495,0.0000033215404,0.000025906707,0.000021985521,0.000012376377,0.00025272777],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9877244,0.00820878,0.000404037,0.00068376726,0.0024920169,0.0004870926],"domain_scores_gemma":[0.89532864,0.0823031,0.01278811,0.0024328444,0.005766271,0.001381006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063857697,0.00031268824,0.00019161885,0.0006410873,0.00048923976,0.002211566,0.00041660428,0.00086622156,0.0022502986],"category_scores_gemma":[0.064690895,0.00028914513,0.00023236692,0.00045313148,0.0013813818,0.0024055054,0.0007996453,0.00089302385,0.00036037335],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001129077,0.0013629214,0.6959507,0.0010588298,0.00039933514,0.00084957905,0.044905774,0.011658915,0.044403218,0.026956923,0.0022951043,0.16902956],"study_design_scores_gemma":[0.00010088004,0.0011259409,0.9326545,0.00014733736,0.00012685925,0.00032385468,0.020961411,0.013413646,0.007174763,0.01663284,0.0072522713,0.00008569934],"about_ca_topic_score_codex":0.0017000182,"about_ca_topic_score_gemma":0.002317934,"teacher_disagreement_score":0.0063857697,"about_ca_system_score_codex":0.0012820923,"about_ca_system_score_gemma":0.0008657908,"threshold_uncertainty_score":0.033771574},"labels":[],"label_agreement":null},{"id":"W2142259572","doi":"10.1109/cmpsac.1994.342817","title":"Measuring program structure with inter-module metrics","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Modularity (biology); Cohesion (chemistry); Modular design; Programming language; Structured programming; Program structure; Software; Vocabulary; Software engineering; Quality (philosophy); Modular programming; Restructuring; Software quality; Software development","score_opus":0.04112952568894309,"score_gpt":0.2504575050966107,"score_spread":0.2093279794076676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142259572","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25578076,0.00046654098,0.73325294,0.0002766866,0.000031515825,0.00046521323,0.0008274757,0.0042921794,0.0046065687],"genre_scores_gemma":[0.6152912,0.00015004732,0.38044745,0.000036989266,0.000025934814,0.00064847677,0.0020425438,0.0005073954,0.00084991794],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98833627,0.0034053302,0.0011968856,0.00083512993,0.0059001273,0.00032621794],"domain_scores_gemma":[0.94647264,0.021766335,0.010969289,0.007993665,0.011854788,0.0009433279],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006286168,0.00093607005,0.00070474274,0.007028016,0.00064703095,0.0021866236,0.0013118571,0.0009166259,0.0012440343],"category_scores_gemma":[0.049019832,0.00042982673,0.0005138705,0.0052092467,0.0009995474,0.005225799,0.0019873665,0.0008299335,0.00037544913],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038019908,0.00059511577,0.15566157,0.0009879644,0.00042890708,0.00019654745,0.0018383254,0.0973542,0.043434322,0.0362292,0.00422607,0.65866756],"study_design_scores_gemma":[0.00012362594,0.0019783084,0.13335706,0.00024350926,0.00032141016,0.00067949545,0.0007384327,0.6690946,0.105475046,0.06848885,0.019240847,0.0002587603],"about_ca_topic_score_codex":0.0010710867,"about_ca_topic_score_gemma":0.0012427049,"teacher_disagreement_score":0.007028016,"about_ca_system_score_codex":0.0012095505,"about_ca_system_score_gemma":0.0010239795,"threshold_uncertainty_score":0.03324485},"labels":[],"label_agreement":null},{"id":"W2142478499","doi":"10.1109/icpc.2009.5090073","title":"Improving program comprehension by enhancing program constructs: An analysis of the Umple language","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Program comprehension; Computer science; Programming language; Syntax; Unified Modeling Language; Comprehension; Program analysis; Abstract syntax tree; Set (abstract data type); Object-oriented programming; Software engineering; Natural language processing; Software; Software system","score_opus":0.007793741832048365,"score_gpt":0.28994855853316587,"score_spread":0.28215481670111753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142478499","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2675985,0.00050479,0.71511,0.0007918969,0.000027170401,0.00016368995,0.00009617304,0.0046031377,0.011104707],"genre_scores_gemma":[0.6980982,0.00034366094,0.2953598,0.00039595686,0.000050718158,0.00013731369,0.00015034333,0.0015428595,0.0039212187],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9983802,0.00057520287,0.00006014422,0.00017664039,0.000585916,0.00022194644],"domain_scores_gemma":[0.9909451,0.005416483,0.0011283264,0.0013881272,0.0009699418,0.0001519522],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019149046,0.0006837229,0.00042925266,0.0009356626,0.0005380697,0.0012286877,0.0011859925,0.00074739166,0.0022611162],"category_scores_gemma":[0.010194475,0.00048670868,0.0009923852,0.00062757195,0.0020396323,0.0056893695,0.0019097432,0.0016469363,0.0002993708],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010889594,0.0008599794,0.018720748,0.0013802584,0.00013653192,0.001834141,0.0054141977,0.088929035,0.14347337,0.40234274,0.0063930927,0.3294269],"study_design_scores_gemma":[0.00016619528,0.0011228971,0.008937145,0.00029345774,0.00034537565,0.0016209135,0.0008908727,0.63445336,0.16285862,0.1574699,0.031695656,0.0001455553],"about_ca_topic_score_codex":0.000581943,"about_ca_topic_score_gemma":0.0007715818,"teacher_disagreement_score":0.0022611162,"about_ca_system_score_codex":0.00058963266,"about_ca_system_score_gemma":0.0008888914,"threshold_uncertainty_score":0.010127127},"labels":[],"label_agreement":null},{"id":"W2142716611","doi":"10.1145/1137983.1138013","title":"Information theoretic evaluation of change prediction models for large-scale software","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Zipf's law; Computer science; Closeness; Data mining; Probabilistic logic; Software; Entropy (arrow of time); Principle of maximum entropy; Algorithm; Mathematics; Statistics; Artificial intelligence","score_opus":0.04478003771689125,"score_gpt":0.2831079990688019,"score_spread":0.23832796135191064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142716611","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49836907,0.0024792429,0.49174514,0.0020819977,0.000091220194,0.00023535526,0.0007523927,0.0012873935,0.0029581331],"genre_scores_gemma":[0.95997113,0.0003335866,0.038425192,0.00014326871,0.00007761414,0.00009303115,0.00066370907,0.00004557851,0.0002469809],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.992074,0.004144573,0.0004130449,0.00087133725,0.0021963727,0.00030059],"domain_scores_gemma":[0.82773423,0.15520959,0.0064192163,0.0045377933,0.004993894,0.0011053412],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.02582177,0.0015432686,0.0014936499,0.0039662016,0.0010327386,0.002069824,0.002124312,0.0020603968,0.0008002414],"category_scores_gemma":[0.08743717,0.0007064696,0.0012186117,0.0022504837,0.0022194837,0.00539952,0.0019253927,0.0019235901,0.00016983498],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031406988,0.00013041338,0.007933081,0.00007709579,0.00017216323,0.00005455312,0.00009025996,0.9683796,0.0003175851,0.0064026765,0.00038077118,0.015747631],"study_design_scores_gemma":[0.000007157922,0.000046072244,0.0008642592,0.000006920462,0.000012152034,0.00001326883,0.000011259178,0.9958091,0.00021132483,0.0029708375,0.000037096564,0.000010522547],"about_ca_topic_score_codex":0.008292457,"about_ca_topic_score_gemma":0.0055169733,"teacher_disagreement_score":0.97417825,"about_ca_system_score_codex":0.005556827,"about_ca_system_score_gemma":0.0018813628,"threshold_uncertainty_score":0.13656026},"labels":[],"label_agreement":null},{"id":"W2142782287","doi":"10.1109/wcre.2005.31","title":"Symbolic Interpretation of Legacy Assembly Language","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Programming language; Computer science; Abstract interpretation; IBM; Symbolic execution; Semantics (computer science); Symbolic data analysis; Interpretation (philosophy); Program analysis; Operational semantics; Symbolic trajectory evaluation; Code (set theory); Control flow; Static program analysis; The Symbolic; Theoretical computer science; Model checking; Software; Software development","score_opus":0.00632289803128751,"score_gpt":0.2603338361779754,"score_spread":0.2540109381466879,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142782287","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046498884,0.00022957855,0.9250407,0.00036361616,0.00006419606,0.000042240918,0.00025895302,0.0050401688,0.02246172],"genre_scores_gemma":[0.67806387,0.0005153112,0.309232,0.00013150062,0.000061257255,0.00007841821,0.0005842574,0.0011983629,0.010134961],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.999411,0.00013987311,0.000033895492,0.00006598688,0.0002774545,0.0000717196],"domain_scores_gemma":[0.9991949,0.0002464999,0.00010273773,0.00019741913,0.0002427315,0.000015580328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00048620603,0.00072238606,0.00032438175,0.0011584813,0.00059883425,0.0015331251,0.0009193807,0.0004257575,0.0040612714],"category_scores_gemma":[0.0022382617,0.00029188162,0.0005506987,0.0007707708,0.002069048,0.0012091789,0.0008653339,0.0008180185,0.0006080116],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012707374,0.000052165433,0.0011582556,0.00033219534,0.00002728454,0.0017513739,0.0019881583,0.12450043,0.026161915,0.724678,0.003724318,0.11549879],"study_design_scores_gemma":[0.00004509479,0.00008344932,0.00056811055,0.00015943548,0.00006820204,0.0005132324,0.00034530315,0.49009687,0.056100335,0.39778206,0.05416726,0.00007055549],"about_ca_topic_score_codex":0.0030044145,"about_ca_topic_score_gemma":0.0036145533,"teacher_disagreement_score":0.0040612714,"about_ca_system_score_codex":0.0011616893,"about_ca_system_score_gemma":0.0009929385,"threshold_uncertainty_score":0.013586342},"labels":[],"label_agreement":null},{"id":"W2142920827","doi":"10.1007/s10664-015-9398-0","title":"Introduction to the special issue on software maintenance and evolution research","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Software development; Open-source software development; Software engineering; Software analytics; Computer science; World Wide Web; Software; Source code; Open source software; Engineering; Software development process; Operating system","score_opus":0.04403391377686007,"score_gpt":0.31887500522402334,"score_spread":0.2748410914471633,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142920827","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00035034976,0.073383965,0.006056894,0.044122063,0.85010237,0.00007683256,0.00081512117,0.00036651152,0.024726002],"genre_scores_gemma":[0.0013895561,0.03429588,0.0018928244,0.01565363,0.8825177,0.000080790116,0.0011702651,0.00049049634,0.06250876],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975873,0.00039879896,0.00028974365,0.00042968095,0.0010832191,0.0002112722],"domain_scores_gemma":[0.98200315,0.008157221,0.0011253036,0.0012329264,0.0046291985,0.0028521041],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037286298,0.0025504455,0.0035754424,0.008578564,0.0016639109,0.008407955,0.002466691,0.0042861826,0.0981949],"category_scores_gemma":[0.013129992,0.0008513262,0.0020331086,0.005620956,0.0016352867,0.007657638,0.0035198778,0.007751443,0.04855593],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000016634018,0.000031546948,0.00010361124,0.00023326873,0.000011552282,0.00003423935,0.00001795817,0.00006909893,0.00013697476,0.0017490452,0.9697807,0.0278154],"study_design_scores_gemma":[0.000009301798,0.00004138888,0.0006306006,0.00033121384,0.000016731585,0.00015298845,0.000033398654,0.0001612958,0.00006867837,0.004931184,0.9936067,0.000016476772],"about_ca_topic_score_codex":0.0009198494,"about_ca_topic_score_gemma":0.002191762,"teacher_disagreement_score":0.0981949,"about_ca_system_score_codex":0.0016744925,"about_ca_system_score_gemma":0.0023594429,"threshold_uncertainty_score":0.32849467},"labels":[],"label_agreement":null},{"id":"W2142971208","doi":"10.1145/1117696.1117711","title":"NaCIN","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Task (project management); Eclipse; Process (computing); Code (set theory); Software engineering; Programming language; Engineering drawing; Systems engineering; Engineering; Set (abstract data type)","score_opus":0.01427911132402806,"score_gpt":0.2620745987438099,"score_spread":0.24779548741978183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142971208","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039283402,0.0023791387,0.35965818,0.002089515,0.0014468115,0.0010490732,0.012664117,0.28095415,0.3004756],"genre_scores_gemma":[0.20890553,0.0021406736,0.41742334,0.0037291632,0.00032215513,0.0010868792,0.0575276,0.04412889,0.2647358],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981926,0.00028301234,0.000106690604,0.00037659638,0.0008997201,0.0001412863],"domain_scores_gemma":[0.99641114,0.0007561213,0.00014367048,0.0011639668,0.0013075516,0.0002175411],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019374895,0.0007748459,0.00047773522,0.0010043293,0.00064885576,0.0022222847,0.001646516,0.0006750466,0.026876297],"category_scores_gemma":[0.0057853158,0.0004904869,0.00041114533,0.0006357242,0.00037722068,0.0025891894,0.002197688,0.0010080179,0.017655665],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010319824,0.0002479525,0.0080604255,0.0009699914,0.000051432227,0.00036414954,0.00071984436,0.0022576647,0.026881099,0.03185436,0.29842588,0.6291353],"study_design_scores_gemma":[0.00003999508,0.000114872695,0.0031360665,0.00010300039,0.000028132596,0.00050612725,0.000098173965,0.0074251574,0.01033397,0.004153506,0.9740172,0.000043708376],"about_ca_topic_score_codex":0.0035408814,"about_ca_topic_score_gemma":0.008296379,"teacher_disagreement_score":0.026876297,"about_ca_system_score_codex":0.00090072944,"about_ca_system_score_gemma":0.0014773358,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2142978025","doi":"10.1145/1858996.1859015","title":"Deviance from perfection is a better criterion than closeness to evil when identifying risky code","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Perfection; Code (set theory); Computer science; Closeness; Detector; Normality; Deviance (statistics); Measure (data warehouse); Set (abstract data type); Robustness (evolution); Algorithm; Artificial intelligence; Theoretical computer science; Computer security; Machine learning; Programming language; Data mining; Mathematics; Statistics; Epistemology","score_opus":0.02695663459996754,"score_gpt":0.2989857409719689,"score_spread":0.27202910637200134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142978025","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73220444,0.00025077077,0.2613814,0.0005077394,0.000047856098,0.00012869423,0.00035489368,0.00092808687,0.004196183],"genre_scores_gemma":[0.957006,0.000035343666,0.042238064,0.0000491127,0.000015868589,0.000046818815,0.00018666634,0.00006688385,0.0003552244],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9915834,0.0021760166,0.00088322564,0.0015667661,0.0033503058,0.00044033644],"domain_scores_gemma":[0.8978897,0.06538336,0.016764212,0.00871353,0.008368377,0.002880859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007900595,0.0009159964,0.0012960266,0.0063833403,0.00088504306,0.0026960236,0.00094503904,0.0016350081,0.0012778267],"category_scores_gemma":[0.06458483,0.000383389,0.00087815284,0.0021263568,0.0033020189,0.0034228081,0.0022764446,0.002554116,0.0004130018],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011170874,0.0006134373,0.55997825,0.00046751418,0.00064781355,0.0005774213,0.0022427524,0.13022561,0.02516838,0.031388693,0.0021758552,0.24539717],"study_design_scores_gemma":[0.00012707453,0.0027219374,0.24545555,0.00018975283,0.00021783297,0.0018734559,0.001271665,0.6063559,0.044893686,0.091857195,0.004466389,0.0005694863],"about_ca_topic_score_codex":0.0011736628,"about_ca_topic_score_gemma":0.0016244258,"teacher_disagreement_score":0.007900595,"about_ca_system_score_codex":0.0007891666,"about_ca_system_score_gemma":0.00094794057,"threshold_uncertainty_score":0.041782856},"labels":[],"label_agreement":null},{"id":"W2143040155","doi":"10.1145/1882362.1882435","title":"The impact of social media on software engineering practices and tools","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":215,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software engineering; Social media; Social software engineering; Software; Software development; Software construction; World Wide Web; Programming language","score_opus":0.026812702424623455,"score_gpt":0.31492531766555576,"score_spread":0.2881126152409323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143040155","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89968026,0.0031563798,0.010095325,0.010618849,0.00020523222,0.00008573933,0.00014183721,0.00018038401,0.07583595],"genre_scores_gemma":[0.99496967,0.0011934101,0.00199852,0.0002581271,0.00022131401,0.00003920113,0.00003171512,0.00005530682,0.0012326712],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97623277,0.015267278,0.00085991295,0.001362514,0.005232438,0.0010451281],"domain_scores_gemma":[0.776756,0.18276408,0.018541016,0.008451474,0.008950506,0.004537079],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012003306,0.0005310521,0.00032674932,0.0050914492,0.0036385613,0.010795969,0.0010101531,0.0013119214,0.0032768815],"category_scores_gemma":[0.067733064,0.0005898592,0.0004094517,0.0030134649,0.00513644,0.010274167,0.006153351,0.001675298,0.00055423187],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053549267,0.000847483,0.28387567,0.0011086569,0.0004995645,0.001949842,0.09844998,0.0028125134,0.008193358,0.0682144,0.0058119944,0.527701],"study_design_scores_gemma":[0.00017227985,0.0017079095,0.53166467,0.0024129276,0.00076146866,0.003491365,0.16405548,0.017901527,0.012098801,0.119504675,0.14568378,0.0005450375],"about_ca_topic_score_codex":0.00322885,"about_ca_topic_score_gemma":0.005221639,"teacher_disagreement_score":0.012003306,"about_ca_system_score_codex":0.0018973869,"about_ca_system_score_gemma":0.0016724381,"threshold_uncertainty_score":0.06348032},"labels":[],"label_agreement":null},{"id":"W2143326878","doi":"10.1109/wpc.2004.1311050","title":"Giving meaning to macros","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code refactoring; Computer science; Preprocessor; Programming language; Readability; Maintainability; Code (set theory); Macro; Software engineering; Set (abstract data type); Source code; Software maintenance; Software; Software system","score_opus":0.01557246605876644,"score_gpt":0.2683224698690844,"score_spread":0.252750003810318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143326878","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034033667,0.00401432,0.83745307,0.005882779,0.0032467782,0.0002110086,0.0035398558,0.021887098,0.08973142],"genre_scores_gemma":[0.41254026,0.0048660003,0.5069844,0.003864511,0.0013685267,0.00031697645,0.0030246407,0.018903468,0.048131235],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978667,0.00068707886,0.00022694696,0.00046967246,0.00050962897,0.00024000663],"domain_scores_gemma":[0.9938408,0.0025711616,0.00040371873,0.0021632903,0.00079935783,0.00022153946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024267929,0.0009639288,0.00079081795,0.0017346515,0.0020319498,0.0056497063,0.0011922205,0.0013967986,0.013820689],"category_scores_gemma":[0.011933089,0.0009946912,0.0010025324,0.0017395514,0.00460576,0.008269482,0.00443346,0.0024627775,0.0048670755],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023858955,0.000028749533,0.0042367834,0.0003942134,0.000043663564,0.0008287988,0.0056571257,0.0016946185,0.009703242,0.7786977,0.04425084,0.15422563],"study_design_scores_gemma":[0.000032659278,0.000062663275,0.0022807345,0.00024038728,0.00008369574,0.00048498198,0.00084342,0.0041793827,0.010353596,0.344032,0.637327,0.00007949109],"about_ca_topic_score_codex":0.0021386764,"about_ca_topic_score_gemma":0.0033442338,"teacher_disagreement_score":0.013820689,"about_ca_system_score_codex":0.0014008698,"about_ca_system_score_gemma":0.0012581792,"threshold_uncertainty_score":0.046234787},"labels":[],"label_agreement":null},{"id":"W2143508381","doi":"10.1109/coginf.2003.1225957","title":"Adding diversity to software inspections","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Process (computing); Computer science; Diversity (politics); Software; Selection (genetic algorithm); Software engineering; Software inspection; Software development process; Mechanism (biology); Software development; Process management; Data science; Software quality; Artificial intelligence; Engineering; Programming language","score_opus":0.022478801506589464,"score_gpt":0.2564174281779501,"score_spread":0.23393862667136064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143508381","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21719499,0.0047077416,0.74101084,0.0055754203,0.00035816178,0.00032594023,0.00012218913,0.0011291992,0.029575562],"genre_scores_gemma":[0.84029377,0.0010920928,0.15466908,0.00066779787,0.00052594586,0.00017912197,0.00011207756,0.00012567997,0.002334434],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9779519,0.009719903,0.0009075659,0.0027450284,0.0073171128,0.0013584879],"domain_scores_gemma":[0.920653,0.04970013,0.006326414,0.012640893,0.006898496,0.003781062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013297379,0.00093924254,0.0015413675,0.0046079447,0.0027150773,0.0043115118,0.0026127258,0.0022506467,0.0017353691],"category_scores_gemma":[0.06572936,0.00085648155,0.0013817212,0.002168934,0.004607973,0.0067101945,0.012155566,0.0034761124,0.00055559137],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006856804,0.00044792923,0.02934705,0.00059728615,0.00044218072,0.00077800977,0.0062869843,0.08841193,0.012360318,0.15771107,0.0035880622,0.69934344],"study_design_scores_gemma":[0.0002953675,0.0012142985,0.014538238,0.00036435746,0.0003369804,0.0017632273,0.0015652068,0.14782007,0.0101261735,0.7869106,0.034767862,0.0002976219],"about_ca_topic_score_codex":0.00077529193,"about_ca_topic_score_gemma":0.0011281047,"teacher_disagreement_score":0.013297379,"about_ca_system_score_codex":0.0019452141,"about_ca_system_score_gemma":0.001827422,"threshold_uncertainty_score":0.07032412},"labels":[],"label_agreement":null},{"id":"W2143568389","doi":"10.1109/icpc.2008.19","title":"Identifying Architectural Change Patterns in Object-Oriented Systems","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Architectural pattern; Software evolution; Object-oriented programming; Software system; Software architecture; Software engineering; Object (grammar); Software; Architecture; Software maintenance; Programming language; Artificial intelligence; Software construction","score_opus":0.06373804501398189,"score_gpt":0.283282740943843,"score_spread":0.2195446959298611,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143568389","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91517097,0.00041147528,0.08146043,0.00025743083,0.000016227528,0.00030488154,0.00020730609,0.0006707046,0.0015005021],"genre_scores_gemma":[0.9270368,0.00015555798,0.071639836,0.000049813385,0.0000087430835,0.00011586798,0.00039223395,0.00006186324,0.00053919654],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974733,0.00066050375,0.0003005835,0.0005077654,0.0008725695,0.0001852387],"domain_scores_gemma":[0.9878364,0.004401451,0.0039012118,0.0018785389,0.0016396482,0.00034266064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015478956,0.00038015025,0.0004905821,0.004959075,0.0009915044,0.001792179,0.0006983149,0.0011403859,0.00041624028],"category_scores_gemma":[0.012462141,0.00045659544,0.0004550061,0.0039781732,0.0011739632,0.002370369,0.0011339108,0.0006880892,0.000115688556],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045985743,0.0005873796,0.6068924,0.00054004765,0.000184354,0.0026094872,0.019439684,0.026613561,0.03411797,0.014033979,0.0012798266,0.29324135],"study_design_scores_gemma":[0.000087345594,0.00070588687,0.5933101,0.00020884849,0.00026441694,0.0034225325,0.0118465265,0.30957797,0.022693006,0.042560883,0.015128428,0.00019403767],"about_ca_topic_score_codex":0.0069521484,"about_ca_topic_score_gemma":0.011464927,"teacher_disagreement_score":0.0069521484,"about_ca_system_score_codex":0.001134671,"about_ca_system_score_gemma":0.000657086,"threshold_uncertainty_score":0.01382333},"labels":[],"label_agreement":null},{"id":"W2143654857","doi":"10.1049/iet-sen.2008.0078","title":"Approach for solving the feature location problem by measuring the component modification impact","year":2009,"lang":"en","type":"article","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Feature (linguistics); Component (thermodynamics); Feature model; Relevance (law); Software system; Dependency graph; Data mining; Ranking (information retrieval); Task (project management); Software; Domain (mathematical analysis); Source code; Dependency (UML); Artificial intelligence; Graph; Machine learning; Engineering; Systems engineering; Programming language; Theoretical computer science","score_opus":0.028858509463813898,"score_gpt":0.2773965040843011,"score_spread":0.2485379946204872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143654857","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03308368,0.00023542347,0.96199834,0.00016854361,0.000039902985,0.00016462566,0.00022510374,0.0023407475,0.0017437171],"genre_scores_gemma":[0.32713038,0.00016299794,0.67078793,0.000077857585,0.000037336387,0.0001892857,0.0003066874,0.00013699883,0.001170628],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964804,0.0007713281,0.00024164192,0.0007856705,0.0015768637,0.00014419632],"domain_scores_gemma":[0.9907871,0.003630837,0.001980153,0.0015292757,0.0018910834,0.0001815304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023067421,0.0022485391,0.0013820825,0.007121995,0.00066331885,0.0016018503,0.002205331,0.0022056052,0.0018936575],"category_scores_gemma":[0.009441898,0.00070580066,0.00095026987,0.0032758494,0.0008103313,0.002103409,0.0012066282,0.0015373396,0.0010449885],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042225106,0.0011870048,0.0430547,0.0007510866,0.0008450072,0.0003388934,0.00038279855,0.095128626,0.10660747,0.0108625945,0.0035997601,0.7368198],"study_design_scores_gemma":[0.00009211194,0.0010761422,0.022542413,0.000063327054,0.00043859993,0.0010527357,0.00023962323,0.8995099,0.056260634,0.012868348,0.005651809,0.00020445226],"about_ca_topic_score_codex":0.0032418931,"about_ca_topic_score_gemma":0.0039045543,"teacher_disagreement_score":0.007121995,"about_ca_system_score_codex":0.0011118804,"about_ca_system_score_gemma":0.0013388047,"threshold_uncertainty_score":0.012199402},"labels":[],"label_agreement":null},{"id":"W2143728276","doi":"10.1109/icsm.2004.1357812","title":"Predicting change propagation in software systems","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":230,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Heuristics; Software system; Software evolution; Software; Source code; Software development; Open source software; Software engineering; Software metric; Software quality; Software construction; Programming language; Operating system","score_opus":0.03377002105423018,"score_gpt":0.26066487273447836,"score_spread":0.22689485168024817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143728276","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9578171,0.00059174816,0.03785596,0.00021871601,0.000033424447,0.00016015607,0.0014705361,0.0010001863,0.0008522677],"genre_scores_gemma":[0.96744657,0.00018822616,0.02949961,0.000042209496,0.000029901288,0.000061315084,0.002466613,0.000040966836,0.00022470132],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965061,0.0010914599,0.00037732613,0.0008570582,0.000887872,0.00028026756],"domain_scores_gemma":[0.91253704,0.0667201,0.010746182,0.002889558,0.0058440357,0.0012630348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046254173,0.0010335254,0.0007153176,0.008195108,0.0007404693,0.00173673,0.0008919771,0.00201372,0.0005671886],"category_scores_gemma":[0.046401322,0.00063190644,0.00075297203,0.004965403,0.0006954992,0.002757212,0.0007541933,0.0014859933,0.00038748427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003451024,0.00032237166,0.70263636,0.00024033619,0.00025449463,0.00018835826,0.0005466195,0.18683302,0.0021369522,0.00079996244,0.0017072483,0.10398926],"study_design_scores_gemma":[0.0000406775,0.00031796974,0.16847515,0.000029659599,0.000108390304,0.00024382563,0.00023928503,0.82398087,0.0034558834,0.0022540789,0.0008002263,0.00005406498],"about_ca_topic_score_codex":0.017968321,"about_ca_topic_score_gemma":0.018236937,"teacher_disagreement_score":0.017968321,"about_ca_system_score_codex":0.0012571005,"about_ca_system_score_gemma":0.0010527435,"threshold_uncertainty_score":0.0357275},"labels":[],"label_agreement":null},{"id":"W2143785121","doi":"10.1109/wcre.2008.46","title":"Towards a Process for Developing Maintenance Tools in Academia","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Process (computing); Computer science; Process management; Software engineering; Systems engineering; Engineering; Programming language","score_opus":0.0942352701721534,"score_gpt":0.3467760123492064,"score_spread":0.252540742177053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143785121","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009904601,0.00044564583,0.97826594,0.0035030956,0.00006182674,0.00089808775,0.000023382963,0.0006898872,0.0062074494],"genre_scores_gemma":[0.0459056,0.00047509218,0.95002514,0.00025748735,0.000043255928,0.00072962884,0.000101658945,0.00014987112,0.0023122404],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9435638,0.029570024,0.0051385523,0.004744321,0.015360245,0.0016229269],"domain_scores_gemma":[0.9128163,0.038405586,0.00831238,0.0183643,0.018171325,0.0039302027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07395619,0.0017432151,0.001160466,0.006320776,0.0067036445,0.018775618,0.00514939,0.008339062,0.0021532788],"category_scores_gemma":[0.073526144,0.0022724497,0.0017412199,0.0043179872,0.010898744,0.019046636,0.010325979,0.0076318453,0.0028685236],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010211405,0.00069445895,0.003924051,0.0009856002,0.000045041277,0.0009915939,0.035053138,0.008365259,0.012144514,0.5526658,0.005221562,0.37980688],"study_design_scores_gemma":[0.0002292982,0.0015069527,0.003818096,0.0024220995,0.00014488699,0.0023241332,0.017900893,0.05756821,0.03378893,0.528608,0.35129094,0.00039753347],"about_ca_topic_score_codex":0.0022933863,"about_ca_topic_score_gemma":0.0018739812,"teacher_disagreement_score":0.07395619,"about_ca_system_score_codex":0.005348109,"about_ca_system_score_gemma":0.022329673,"threshold_uncertainty_score":0.3911224},"labels":[],"label_agreement":null},{"id":"W2143818143","doi":"10.1145/1985441.1985467","title":"Modeling the evolution of topics in source code histories","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Source code; Software evolution; Code (set theory); Programming language; Software development; Software; Software construction","score_opus":0.041987563307647385,"score_gpt":0.24553846296164972,"score_spread":0.20355089965400233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143818143","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5890855,0.0009361012,0.40496972,0.00068296754,0.00004140609,0.00019869521,0.0013033596,0.0008993087,0.0018828975],"genre_scores_gemma":[0.9247726,0.0004971421,0.0693471,0.00006646393,0.000060488168,0.0003094023,0.00209463,0.00014701886,0.0027050022],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99863845,0.00045238127,0.00009045706,0.00050531747,0.0002002442,0.00011311608],"domain_scores_gemma":[0.9831872,0.0130132865,0.0016239427,0.00091739424,0.0009022666,0.00035591546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038092665,0.0007245906,0.00059509446,0.004054355,0.0005455644,0.0017681598,0.001229424,0.0014888587,0.00091566175],"category_scores_gemma":[0.022142489,0.0007791199,0.0010368298,0.0024586562,0.000775473,0.0035547514,0.0011856576,0.0009936298,0.00040476365],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065660186,0.0002682361,0.14388272,0.00045274707,0.0002956487,0.0005861296,0.004595741,0.61091673,0.012653145,0.04093282,0.0036638624,0.18109564],"study_design_scores_gemma":[0.000013942997,0.000030588693,0.005927704,0.000010140564,0.000028541368,0.00007406471,0.000076280645,0.9834879,0.0006869889,0.008702219,0.0009471468,0.00001444636],"about_ca_topic_score_codex":0.011578934,"about_ca_topic_score_gemma":0.009488339,"teacher_disagreement_score":0.011578934,"about_ca_system_score_codex":0.0013183138,"about_ca_system_score_gemma":0.0008340592,"threshold_uncertainty_score":0.023023129},"labels":[],"label_agreement":null},{"id":"W2144128087","doi":"10.1109/rev.2006.8","title":"Visualizing a Requirements-centred Social Network to Maintain Awareness Within Development Teams","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Requirements analysis; Rework; Computer science; Social network (sociolinguistics); Requirements management; Use Case Diagram; Plan (archaeology); Work (physics); Knowledge management; User requirements document; Functional requirement; Process management; Software; Engineering; World Wide Web; Social media; Software engineering; Unified Modeling Language","score_opus":0.03321915956230581,"score_gpt":0.3055645249894578,"score_spread":0.272345365427152,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144128087","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.111991435,0.000491192,0.8641806,0.0022885173,0.00015706985,0.0002651804,0.0007399694,0.0037210376,0.016164973],"genre_scores_gemma":[0.56615794,0.0004605094,0.4285719,0.0001240828,0.000042816853,0.00022566241,0.0006465605,0.00031536698,0.003455178],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.999257,0.00050668686,0.00003094231,0.0000674777,0.00010588896,0.000031890442],"domain_scores_gemma":[0.997029,0.0018951081,0.00027521417,0.00026807567,0.00032824368,0.0002043107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013440734,0.00059594837,0.00021884927,0.0025068806,0.00068177056,0.001668868,0.00057574105,0.00078804995,0.0041016117],"category_scores_gemma":[0.0057246364,0.00030761518,0.00039002264,0.001322739,0.0005232054,0.0035195467,0.0016617522,0.00070978265,0.00042133202],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010630598,0.00041590852,0.02698078,0.0018833294,0.00021588987,0.0024177472,0.076101296,0.124501795,0.04731774,0.20201051,0.030837499,0.4862545],"study_design_scores_gemma":[0.00027145186,0.0003467088,0.01972813,0.0006169416,0.00019886585,0.0016172006,0.016646214,0.58066547,0.015934194,0.13672972,0.22701462,0.00023051072],"about_ca_topic_score_codex":0.004369664,"about_ca_topic_score_gemma":0.004813518,"teacher_disagreement_score":0.004369664,"about_ca_system_score_codex":0.0008111738,"about_ca_system_score_gemma":0.00084948563,"threshold_uncertainty_score":0.013721228},"labels":[],"label_agreement":null},{"id":"W2144282912","doi":"10.1109/wpc.2003.1199206","title":"An optimal algorithm for MoJo distance","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Heuristic; Cluster analysis; Algorithm; Computer science; Measure (data warehouse); Computation; Software; Data mining; Artificial intelligence","score_opus":0.013805915382186987,"score_gpt":0.28465550296322956,"score_spread":0.27084958758104255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144282912","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018316287,0.00020150692,0.9944571,0.00016428692,0.00011200659,0.000062862004,0.000035089226,0.00063299364,0.0025025168],"genre_scores_gemma":[0.038435128,0.00021965225,0.95752025,0.00012648235,0.00011611435,0.00022516288,0.00021527668,0.00029243247,0.002849487],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9964893,0.0007136214,0.00022319626,0.00076920877,0.0014853218,0.00031930133],"domain_scores_gemma":[0.9972638,0.0010674279,0.00014990543,0.00056152575,0.0008089295,0.00014852072],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017936452,0.0011824646,0.0016539061,0.0031310844,0.0019861043,0.0025033215,0.0025597333,0.0025500862,0.008148005],"category_scores_gemma":[0.011843629,0.0007798694,0.0013092957,0.0027216661,0.0016641216,0.0036617909,0.003154342,0.0032927368,0.0032913445],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020191296,0.00018870698,0.00059754134,0.00023220036,0.00006668436,0.000066141,0.00020945987,0.13104285,0.008213413,0.2210873,0.015144947,0.6229488],"study_design_scores_gemma":[0.0000847372,0.00013513758,0.00033485633,0.000054066553,0.000027245522,0.00020972638,0.00009056889,0.7373881,0.004318005,0.23504151,0.022248184,0.000067725654],"about_ca_topic_score_codex":0.0039924677,"about_ca_topic_score_gemma":0.0042473366,"teacher_disagreement_score":0.008148005,"about_ca_system_score_codex":0.002139113,"about_ca_system_score_gemma":0.0029368545,"threshold_uncertainty_score":0.0272578},"labels":[],"label_agreement":null},{"id":"W2144491568","doi":"10.1109/ase.2003.1240336","title":"Detecting requirements interactions: a three-level framework","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Domain (mathematical analysis); Lift (data mining); Formal specification; Requirements analysis; Functional requirement; Software engineering; Distributed computing; Programming language; Data mining; Software; Mathematics","score_opus":0.0848601438171301,"score_gpt":0.3400426612558367,"score_spread":0.2551825174387066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144491568","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002438938,0.00017519906,0.99481404,0.0004097706,0.000012075184,0.00029482952,0.00005347735,0.0009891444,0.0008124399],"genre_scores_gemma":[0.038746424,0.00017183812,0.95969075,0.00015888519,0.000019964229,0.00025353584,0.00027133373,0.00012231486,0.00056496094],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96686393,0.0113888,0.00427167,0.003622117,0.012140884,0.001712602],"domain_scores_gemma":[0.95716846,0.025853826,0.0044843443,0.005113331,0.0060606836,0.0013192727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021977488,0.0030294245,0.0023107312,0.008186714,0.0029978228,0.013248967,0.005399277,0.005648774,0.002913221],"category_scores_gemma":[0.0361977,0.0030997498,0.0066446154,0.0030031223,0.008382867,0.009223862,0.008103597,0.005624243,0.0012699347],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003561273,0.00095625117,0.015735166,0.0033360962,0.0010614075,0.0033877245,0.011317681,0.12757318,0.048978917,0.408518,0.00599957,0.3727799],"study_design_scores_gemma":[0.00016155101,0.0006036179,0.004489688,0.0018338114,0.000727041,0.002364918,0.002380622,0.657972,0.025944078,0.25005308,0.052966308,0.0005032525],"about_ca_topic_score_codex":0.011910947,"about_ca_topic_score_gemma":0.007909106,"teacher_disagreement_score":0.021977488,"about_ca_system_score_codex":0.0039196294,"about_ca_system_score_gemma":0.007646901,"threshold_uncertainty_score":0.116229415},"labels":[],"label_agreement":null},{"id":"W2144528247","doi":"10.1145/1985793.1985813","title":"An empirical study of build maintenance effort","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":139,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Codebase; Source code; Computer science; KPI-driven code analysis; Executable; Code review; Overhead (engineering); Software maintenance; Software engineering; Open source; Code (set theory); Software development; Static program analysis; Software; Operating system; Programming language; Set (abstract data type)","score_opus":0.045975590417711126,"score_gpt":0.32489056661915494,"score_spread":0.2789149762014438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144528247","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9954146,0.00023200872,0.0014812239,0.00011876523,0.0000055729347,0.000037861348,0.0003866272,0.00004046977,0.0022828176],"genre_scores_gemma":[0.99789375,0.00011980445,0.0008441512,0.000028372184,0.000009167493,0.00004232558,0.0005941824,0.000021190166,0.0004469863],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9881402,0.0043343445,0.0014659914,0.0011412918,0.0043321857,0.0005860449],"domain_scores_gemma":[0.535477,0.3329196,0.08894076,0.014597385,0.023243094,0.0048221764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011727772,0.00032030288,0.00023781207,0.0038378558,0.0004721572,0.0012516082,0.00096347247,0.00062572793,0.0017580136],"category_scores_gemma":[0.16963878,0.00042417293,0.00026966317,0.0035854094,0.0010155303,0.0033141137,0.0012221186,0.0011254916,0.0004894615],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012620358,0.00020229704,0.9784184,0.000097211596,0.00008853724,0.00010518801,0.0018586967,0.0007245009,0.00041210075,0.00039954158,0.00056160684,0.017005717],"study_design_scores_gemma":[0.000008756999,0.00030686476,0.99255145,0.000040434956,0.000023189657,0.00032457654,0.0016385142,0.0024979033,0.0005080464,0.0002103296,0.0018756666,0.00001428766],"about_ca_topic_score_codex":0.0015649803,"about_ca_topic_score_gemma":0.001921386,"teacher_disagreement_score":0.011727772,"about_ca_system_score_codex":0.00070026587,"about_ca_system_score_gemma":0.0005845251,"threshold_uncertainty_score":0.062023103},"labels":[],"label_agreement":null},{"id":"W2144641356","doi":"10.1007/s10664-006-7552-4","title":"A flexible method for software effort estimation by analogy","year":2006,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":156,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Analogy; Data mining; Computer science; Similarity (geometry); Quality (philosophy); Feature (linguistics); Adaptation (eye); Estimation; Machine learning; Artificial intelligence; Engineering; Image (mathematics)","score_opus":0.017182083426203802,"score_gpt":0.31183334560405335,"score_spread":0.29465126217784954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144641356","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011775229,0.000030620213,0.9976179,0.000026479755,0.000014770622,0.00002584342,0.000016133996,0.00021433644,0.0008763423],"genre_scores_gemma":[0.11688241,0.00012932476,0.8787065,0.00008074004,0.0000821847,0.00042966942,0.000112616035,0.00018940384,0.0033872125],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9929825,0.0032782767,0.00027478219,0.0012818011,0.0019317886,0.00025088008],"domain_scores_gemma":[0.9895914,0.0061598634,0.00044228247,0.0025484487,0.0011190297,0.00013899474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004920299,0.0011707424,0.001786527,0.002897442,0.0009818826,0.0017615579,0.0030542733,0.0020091352,0.0066309534],"category_scores_gemma":[0.036008995,0.00081021246,0.0016671956,0.0033140557,0.0012161754,0.0048909215,0.00394892,0.0030213988,0.0021221086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014889764,0.0001929295,0.0024013256,0.00019595498,0.00014489466,0.00019417822,0.00038189662,0.07221169,0.0067206277,0.31727135,0.0030785876,0.5970577],"study_design_scores_gemma":[0.00006249863,0.00019388471,0.0018631076,0.00006836536,0.00006941561,0.0004041722,0.000061976905,0.71687174,0.0039101006,0.26755834,0.008845812,0.00009059444],"about_ca_topic_score_codex":0.0011324313,"about_ca_topic_score_gemma":0.0009013037,"teacher_disagreement_score":0.0066309534,"about_ca_system_score_codex":0.0006541079,"about_ca_system_score_gemma":0.00097640615,"threshold_uncertainty_score":0.026021361},"labels":[],"label_agreement":null},{"id":"W2144859782","doi":"10.1109/vlhcc.2008.4639052","title":"Exploring the evolution of software quality with animated visualization","year":2008,"lang":"en","type":"article","venue":"Proceedings/Proceedings -- IEEE Symposium on Visual Languages and Human-Centric Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software visualization; Visualization; Animation; Software quality; Software; Software evolution; Intuition; Quality (philosophy); Software analytics; Task (project management); Visual analytics; Software engineering; Software metric; Software development; Human–computer interaction; Software construction; Artificial intelligence; Programming language; Computer graphics (images); Systems engineering","score_opus":0.043143540400998215,"score_gpt":0.31506910513040104,"score_spread":0.2719255647294028,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144859782","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38973662,0.00029527713,0.602109,0.00051331415,0.00002247251,0.00019301029,0.0003317667,0.0037815522,0.003017],"genre_scores_gemma":[0.7009442,0.0001538768,0.2977635,0.00004264825,0.00001026291,0.00010401691,0.0002525427,0.00021882402,0.0005100666],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990393,0.00052747957,0.00004020546,0.00012803619,0.00019223183,0.0000726939],"domain_scores_gemma":[0.99230075,0.005914283,0.00049042894,0.0006891143,0.00043527147,0.00017003425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020631459,0.0006106156,0.00043086338,0.0018098458,0.00040743774,0.0019006297,0.0010342607,0.000999845,0.0021494667],"category_scores_gemma":[0.010592219,0.00047834357,0.00060083036,0.0007902959,0.000874224,0.001643803,0.001776416,0.0010137386,0.00021496131],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013661894,0.00078702345,0.03646575,0.0010425134,0.00031827384,0.0024348376,0.021981215,0.2756247,0.2511677,0.033717807,0.0047269394,0.37036705],"study_design_scores_gemma":[0.00016981401,0.00058132794,0.023555048,0.00016369368,0.00012756113,0.000950025,0.0017221813,0.89261717,0.044330444,0.025501477,0.010127518,0.00015381262],"about_ca_topic_score_codex":0.0017849853,"about_ca_topic_score_gemma":0.0016580913,"teacher_disagreement_score":0.0021494667,"about_ca_system_score_codex":0.0006381685,"about_ca_system_score_gemma":0.00033510532,"threshold_uncertainty_score":0.010911047},"labels":[],"label_agreement":null},{"id":"W2144951547","doi":"10.1109/csmr.2006.26","title":"Evaluating architectural stability using a metric-based approach","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Metric (unit); Computer science; Architectural pattern; Software architecture; Architecture; Architecture tradeoff analysis method; Stability (learning theory); Set (abstract data type); Software engineering; Software; Software metric; Software evolution; Software system; Reference architecture; Source code; Code (set theory); Software architecture description; Programming language; Software construction; Machine learning; Engineering","score_opus":0.11208022887017738,"score_gpt":0.35020052957833286,"score_spread":0.23812030070815549,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144951547","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28536084,0.001052702,0.70373064,0.00031402512,0.000069407404,0.0004419229,0.0009187296,0.0020449378,0.006066814],"genre_scores_gemma":[0.7615605,0.00022665049,0.23617153,0.00003657932,0.000029145565,0.00027605915,0.0010227178,0.00009499018,0.0005817711],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.990932,0.002530415,0.001133825,0.00056457316,0.0046351384,0.00020414029],"domain_scores_gemma":[0.973306,0.010988124,0.0040007904,0.0026980627,0.00843659,0.0005704335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005994352,0.0011642189,0.0010229587,0.011230563,0.0007892547,0.0023853169,0.0014842708,0.0010331111,0.0011588582],"category_scores_gemma":[0.030193722,0.0003607611,0.0006877597,0.005943283,0.00096901355,0.003458013,0.0013850175,0.0006886481,0.00034074186],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008995706,0.00063867454,0.12863265,0.0010122024,0.0010757938,0.00042706646,0.0014870425,0.21832082,0.06698164,0.029853487,0.0028734424,0.5477977],"study_design_scores_gemma":[0.000060441565,0.0013299213,0.05056046,0.0000930008,0.00023316537,0.0005837488,0.0008198265,0.8947231,0.027744545,0.01783176,0.0058058333,0.00021422551],"about_ca_topic_score_codex":0.002674686,"about_ca_topic_score_gemma":0.0025306959,"teacher_disagreement_score":0.011230563,"about_ca_system_score_codex":0.0016702032,"about_ca_system_score_gemma":0.0009454583,"threshold_uncertainty_score":0.031701505},"labels":[],"label_agreement":null},{"id":"W2144961539","doi":"10.1109/icpc.2011.35","title":"The Influence of the Task on Programmer Behaviour","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; McGill University","keywords":"Computer science; Programmer; Session (web analytics); Task (project management); Surprise; TRACE (psycholinguistics); Programming language; Source code; Software development; Software; World Wide Web","score_opus":0.023948104368765773,"score_gpt":0.25383142774394957,"score_spread":0.2298833233751838,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144961539","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9956226,0.0001329925,0.0024242715,0.00006567534,0.000015109417,0.00004140488,0.00008630382,0.000101007994,0.0015106422],"genre_scores_gemma":[0.99730265,0.000054323737,0.0018448655,0.000034574226,0.000019037245,0.000043150616,0.00017965949,0.00007070031,0.00045105332],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9856711,0.0081434315,0.0012413086,0.0017457827,0.0024926085,0.00070582127],"domain_scores_gemma":[0.65155876,0.29294056,0.025469553,0.014238789,0.008842009,0.006950367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008157846,0.0004269152,0.0004581377,0.0015669754,0.0006315527,0.0020568615,0.00070935296,0.00071278214,0.0012784244],"category_scores_gemma":[0.15446594,0.0004440727,0.00043298776,0.0007745731,0.0007691746,0.0011689306,0.0011467219,0.0008856597,0.00048714687],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022883734,0.0005642636,0.88608205,0.0003909021,0.0002704301,0.000355021,0.008182926,0.001764457,0.025518691,0.00021214849,0.00063852494,0.0737322],"study_design_scores_gemma":[0.000024247574,0.00072250987,0.99228,0.000028490767,0.000047640082,0.00034313093,0.001008822,0.0028860813,0.0014019429,0.00021224935,0.0009941775,0.00005063465],"about_ca_topic_score_codex":0.0019528954,"about_ca_topic_score_gemma":0.0026844165,"teacher_disagreement_score":0.008157846,"about_ca_system_score_codex":0.0005064901,"about_ca_system_score_gemma":0.0007204236,"threshold_uncertainty_score":0.043143332},"labels":[],"label_agreement":null},{"id":"W2145091807","doi":"10.1109/issre.1997.630845","title":"Predicting fault-prone modules with case-based reasoning","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nortel (Canada)","funders":"","keywords":"Computer science; Software quality; Linear discriminant analysis; Case-based reasoning; Artificial intelligence; Data mining; Fault (geology); Quality (philosophy); Software; Machine learning; Reliability (semiconductor); Software development","score_opus":0.020812976149890906,"score_gpt":0.23489205583568432,"score_spread":0.2140790796857934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145091807","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5447914,0.0003605952,0.44890878,0.0005781451,0.0000438049,0.00071282027,0.0005488073,0.00073867216,0.003316935],"genre_scores_gemma":[0.8416813,0.00013070309,0.15684858,0.000037115096,0.000018627577,0.00018613096,0.000676999,0.000026189051,0.00039428135],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959765,0.0022049726,0.00030142898,0.00047464579,0.00090356765,0.0001389579],"domain_scores_gemma":[0.96271926,0.029643714,0.0029101598,0.00239881,0.0019747675,0.00035319148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008042259,0.0010554439,0.00087538845,0.0064658388,0.0005724328,0.0023474188,0.0014612484,0.0012677279,0.0020067473],"category_scores_gemma":[0.04709755,0.0005316323,0.0012925207,0.003774847,0.0008309143,0.0026114685,0.0009475021,0.0011583396,0.0004889598],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009917615,0.0014155598,0.16431476,0.00021894177,0.00035291174,0.0006143322,0.0009203457,0.51881355,0.0018563945,0.009182416,0.002766319,0.29855272],"study_design_scores_gemma":[0.000049981285,0.000096468866,0.008981766,0.000025603402,0.000047903155,0.00016202308,0.0001467245,0.9795059,0.0005938159,0.009938207,0.00042421152,0.00002743446],"about_ca_topic_score_codex":0.00806129,"about_ca_topic_score_gemma":0.0075875036,"teacher_disagreement_score":0.00806129,"about_ca_system_score_codex":0.0017052605,"about_ca_system_score_gemma":0.0010818662,"threshold_uncertainty_score":0.042532027},"labels":[],"label_agreement":null},{"id":"W2145209235","doi":"10.1109/icsm.2003.1235428","title":"Impact analysis and change management of UML models","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":183,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Unified Modeling Language; Computer science; Applications of UML; UML tool; Object Constraint Language; Class diagram; Consistency (knowledge bases); Interdependence; Software engineering; Data mining; Programming language; Software; Artificial intelligence","score_opus":0.0427595090332349,"score_gpt":0.3075478474100297,"score_spread":0.26478833837679483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145209235","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08365021,0.0011613028,0.89889926,0.0010724769,0.00011179312,0.0006419907,0.00026637336,0.0021853612,0.012011175],"genre_scores_gemma":[0.647168,0.00078695006,0.3485942,0.00014839497,0.00012866799,0.00052234187,0.0007699239,0.000352313,0.0015292232],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9656907,0.0148284305,0.0018526622,0.0019256758,0.01468357,0.0010191032],"domain_scores_gemma":[0.93942606,0.03847976,0.006644314,0.0071130362,0.00767184,0.0006649813],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016976768,0.0012381119,0.0007940499,0.012805846,0.0015027534,0.0044069486,0.001929291,0.0014222492,0.0024215393],"category_scores_gemma":[0.08007858,0.00087056565,0.0017899674,0.0052921544,0.0021277112,0.006775188,0.0032405143,0.0015429548,0.00029619847],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036728464,0.00032463737,0.03470758,0.0007310644,0.00045236747,0.00083062507,0.0040421677,0.19525374,0.0068767364,0.17577313,0.0033205664,0.57732004],"study_design_scores_gemma":[0.00008060758,0.00027695714,0.016177988,0.0003351348,0.0003756791,0.0004329144,0.0015007859,0.7612453,0.014187239,0.18095417,0.024270048,0.0001632418],"about_ca_topic_score_codex":0.005252629,"about_ca_topic_score_gemma":0.0038642595,"teacher_disagreement_score":0.016976768,"about_ca_system_score_codex":0.0037748986,"about_ca_system_score_gemma":0.0021355036,"threshold_uncertainty_score":0.089782834},"labels":[],"label_agreement":null},{"id":"W2145367887","doi":"10.1109/wcre.2003.1287243","title":"Studying the chaos of code development","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Source code; KPI-driven code analysis; Code (set theory); Software engineering; Software; Software development; Metric (unit); Process (computing); Software evolution; Suite; Code review; CHAOS (operating system); Static program analysis; Programming language; Operating system; Software construction; Engineering","score_opus":0.041310146488413774,"score_gpt":0.27605140812259177,"score_spread":0.234741261634178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145367887","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93570226,0.0017456219,0.049851645,0.001670239,0.00003599408,0.00007758322,0.0002663595,0.00012939055,0.010520869],"genre_scores_gemma":[0.9922232,0.0004291915,0.0064242915,0.000048583122,0.000029881568,0.000048005757,0.0001458479,0.000038287766,0.0006126519],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9979532,0.00077554287,0.00014331598,0.0002770231,0.00063906185,0.00021184851],"domain_scores_gemma":[0.9572982,0.025056586,0.010743199,0.0024875798,0.0029952163,0.0014193538],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.002328231,0.0002304531,0.00031388924,0.0039975015,0.00090742996,0.0020525898,0.00038954164,0.0005152151,0.0009658176],"category_scores_gemma":[0.03434145,0.00029598025,0.00030027435,0.0020507562,0.0024106416,0.003524274,0.0019643893,0.00080695515,0.00017523221],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003895172,0.00015031022,0.6486743,0.0005606796,0.0002745841,0.00048932777,0.011015273,0.06666306,0.009301652,0.11194728,0.00254392,0.14799015],"study_design_scores_gemma":[0.000056695982,0.0005540478,0.52629507,0.00021029032,0.00011691909,0.0007337226,0.003742649,0.21376242,0.005980026,0.22955199,0.01882979,0.00016625304],"about_ca_topic_score_codex":0.0033199955,"about_ca_topic_score_gemma":0.003060027,"teacher_disagreement_score":0.9976718,"about_ca_system_score_codex":0.0016551284,"about_ca_system_score_gemma":0.0013561525,"threshold_uncertainty_score":0.012313008},"labels":[],"label_agreement":null},{"id":"W2145420385","doi":"10.1109/icsm.2006.26","title":"ESDM - A Method for Developing Evolutionary Scenarios for Analysing the Impact of Historical Changes on Architectural Elements","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Software evolution; Architectural pattern; Software; Change impact analysis; Software engineering; Work (physics); Software system; Software construction; Engineering; Programming language","score_opus":0.040594142007186335,"score_gpt":0.35277170513199996,"score_spread":0.3121775631248136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145420385","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0043336414,0.000029548724,0.99017817,0.000092226284,0.000016446185,0.00025008337,0.00062515796,0.0024478212,0.0020269717],"genre_scores_gemma":[0.01969901,0.000035096236,0.9786226,0.0000173919,0.0000039232123,0.00042943042,0.00053910696,0.00020917873,0.00044426258],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99709904,0.0012860451,0.00031054663,0.00037654996,0.0008443095,0.00008349485],"domain_scores_gemma":[0.9857874,0.0103979865,0.00085749896,0.0015247494,0.0012197918,0.00021254394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067283288,0.0021529223,0.0007110958,0.005017882,0.0007978524,0.001954425,0.0016233472,0.0017904165,0.008497769],"category_scores_gemma":[0.023929838,0.001241853,0.0022852535,0.0022835103,0.0010138269,0.0036936216,0.0025483766,0.0017790266,0.001567023],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000494384,0.0004163027,0.018838916,0.0019910112,0.0006838021,0.0013486324,0.0039823665,0.16905132,0.023060918,0.11307382,0.0140292505,0.6530293],"study_design_scores_gemma":[0.00021368804,0.00020718228,0.0028861302,0.00047370463,0.00020281493,0.0012236808,0.0009911035,0.79142374,0.019657902,0.111444935,0.071065806,0.00020934799],"about_ca_topic_score_codex":0.0015272421,"about_ca_topic_score_gemma":0.0032660707,"teacher_disagreement_score":0.008497769,"about_ca_system_score_codex":0.0008927326,"about_ca_system_score_gemma":0.0012064122,"threshold_uncertainty_score":0.035583258},"labels":[],"label_agreement":null},{"id":"W2145512597","doi":"10.1109/tse.2010.91","title":"Work Item Tagging: Communicating Concerns in Collaborative Software Development","year":2010,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Categorization; Software development; Knowledge management; Mechanism (biology); Work (physics); Domain (mathematical analysis); Software; Vocabulary; Empirical research; Open-source software development; World Wide Web; Software engineering; Data science; Human–computer interaction; Artificial intelligence","score_opus":0.018162273126256012,"score_gpt":0.26037683232553865,"score_spread":0.24221455919928264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145512597","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5092361,0.00039908593,0.47395706,0.0015655408,0.00016455137,0.00060682173,0.00009958325,0.0011382158,0.012832951],"genre_scores_gemma":[0.9094877,0.00012704397,0.08746209,0.0002278391,0.00006513938,0.00028009035,0.00015072565,0.00015477427,0.0020445476],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9395561,0.046860512,0.003241386,0.0031017326,0.0060857683,0.0011545164],"domain_scores_gemma":[0.7933839,0.15259016,0.021327868,0.021165388,0.008843099,0.00268963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03132726,0.0011716936,0.0006533076,0.003176606,0.0047116783,0.006580211,0.0024758433,0.0031772205,0.0015950949],"category_scores_gemma":[0.13499509,0.001057364,0.00071278756,0.0023501664,0.004995032,0.011553073,0.008969507,0.0023543986,0.0007920318],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010186497,0.0006728197,0.14736083,0.0011050138,0.00017373284,0.0019869914,0.33620694,0.0032274423,0.031862386,0.032003377,0.003038152,0.44134364],"study_design_scores_gemma":[0.0003619476,0.0034492188,0.20812926,0.0018061654,0.00085651915,0.0088522965,0.23608063,0.09066894,0.0692078,0.18838382,0.19074383,0.0014595567],"about_ca_topic_score_codex":0.0023048664,"about_ca_topic_score_gemma":0.0019266667,"teacher_disagreement_score":0.03132726,"about_ca_system_score_codex":0.0020538666,"about_ca_system_score_gemma":0.0025652093,"threshold_uncertainty_score":0.16567636},"labels":[],"label_agreement":null},{"id":"W2145544369","doi":"10.1109/icse.2007.90","title":"Tracking Code Clones in Evolving Software","year":2007,"lang":"en","type":"article","venue":"Proceedings/Proceedings - International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":216,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Code refactoring; Computer science; Software; Software maintenance; Cloning (programming); Code (set theory); Source code; Tracking (education); Software system; Software development; Software evolution; Programming language; Software engineering; Software construction; Biology; Genetics; Gene","score_opus":0.0352682557164846,"score_gpt":0.29145121113367156,"score_spread":0.25618295541718694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145544369","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37859893,0.00036413566,0.612258,0.00013444654,0.000031875006,0.00034074928,0.00019808288,0.0068352087,0.0012385672],"genre_scores_gemma":[0.4809055,0.00026118624,0.51638526,0.00007709335,0.000011402696,0.00014862102,0.0005206748,0.00041582616,0.001274475],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965888,0.0006863345,0.0002973707,0.0008437446,0.0014694097,0.00011436279],"domain_scores_gemma":[0.9705119,0.014142736,0.0046521896,0.005829419,0.004426105,0.00043763398],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028771316,0.0004321,0.0005810872,0.0024242694,0.0007546167,0.0015472147,0.0012437743,0.0012035198,0.0004907416],"category_scores_gemma":[0.026016511,0.00053966115,0.00041541242,0.0018396254,0.0007979099,0.0020621065,0.0017865101,0.00097485987,0.00024576188],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004122694,0.00031923736,0.051462945,0.00044130444,0.00009219489,0.0009994009,0.0045277295,0.01970185,0.17746891,0.009049096,0.0016577287,0.73386735],"study_design_scores_gemma":[0.00011650392,0.0011411319,0.04326807,0.00021068793,0.0003056911,0.003238011,0.0010550333,0.5144797,0.39612582,0.014776778,0.025065415,0.00021713575],"about_ca_topic_score_codex":0.002323264,"about_ca_topic_score_gemma":0.0025884754,"teacher_disagreement_score":0.0028771316,"about_ca_system_score_codex":0.0006391034,"about_ca_system_score_gemma":0.0009835562,"threshold_uncertainty_score":0.015215933},"labels":[],"label_agreement":null},{"id":"W2145745392","doi":"10.1109/icpc.2007.21","title":"Evaluating Aspect Mining Techniques: A Case Study","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Java; Aspect-oriented programming; Data mining; Data science; Software engineering; Software; Programming language","score_opus":0.09750389719500419,"score_gpt":0.4261712976980123,"score_spread":0.32866740050300813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145745392","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9660406,0.00034324476,0.02878394,0.00050014415,0.000027023532,0.0010326077,0.00013967625,0.00019023889,0.0029424943],"genre_scores_gemma":[0.8513653,0.0006550595,0.14532717,0.000187465,0.00003062793,0.0006233248,0.00025211094,0.00008635706,0.0014726496],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99005353,0.0056972634,0.00065593293,0.0006636494,0.0024345394,0.00049499597],"domain_scores_gemma":[0.9380235,0.049902022,0.0022837932,0.0037091074,0.00525822,0.0008233994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01184823,0.0010144927,0.0008420889,0.002163491,0.002105025,0.0016861872,0.0023289188,0.0031283605,0.0007875417],"category_scores_gemma":[0.029771999,0.00056348037,0.0009987858,0.0023159888,0.0011502584,0.002034789,0.0016479,0.0014089503,0.00029787552],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031056507,0.019532522,0.097898655,0.004120945,0.00057874166,0.026525402,0.041712973,0.06514499,0.090932146,0.008295398,0.0068033156,0.6353492],"study_design_scores_gemma":[0.0027754514,0.031422,0.11642807,0.001460893,0.0016943269,0.02829341,0.04926128,0.44332314,0.24781367,0.012416645,0.064400606,0.0007104715],"about_ca_topic_score_codex":0.002727171,"about_ca_topic_score_gemma":0.006626533,"teacher_disagreement_score":0.01184823,"about_ca_system_score_codex":0.0013667393,"about_ca_system_score_gemma":0.0012521175,"threshold_uncertainty_score":0.06266016},"labels":[],"label_agreement":null},{"id":"W2145766063","doi":"10.1007/s10664-012-9206-z","title":"Introduction to the Special Issue on Mining Software Repositories in 2010","year":2012,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"European Metrology Programme for Innovation and Research; Universität des Saarlandes; University of Calgary","keywords":"Computer science; Software engineering; Data science; Software; Operating system","score_opus":0.018376988051948446,"score_gpt":0.27463512077081664,"score_spread":0.2562581327188682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145766063","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002670291,0.071336776,0.046739213,0.097697064,0.7241182,0.0003141597,0.0076147686,0.0026828025,0.046826784],"genre_scores_gemma":[0.009205637,0.056208592,0.024488289,0.036847804,0.5639704,0.00022859809,0.015011961,0.0024237158,0.291615],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99674964,0.00050208595,0.00043880029,0.0006746676,0.0014424913,0.00019232427],"domain_scores_gemma":[0.97942847,0.005717258,0.001265895,0.0018065083,0.008996169,0.0027857563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004973392,0.0015359275,0.0021767563,0.010011235,0.0015935518,0.007894778,0.0016599202,0.0024976009,0.052198905],"category_scores_gemma":[0.018070837,0.00083187054,0.0012969348,0.0075341803,0.0011126089,0.0077267303,0.003110318,0.0043376265,0.03564903],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003230402,0.00003573131,0.0002864532,0.00023578167,0.000014918198,0.000045093373,0.00002859669,0.0001461371,0.00040996674,0.0012604883,0.92885804,0.06864651],"study_design_scores_gemma":[0.000008965121,0.000050224906,0.00087674183,0.00021410691,0.000016635837,0.0002156467,0.000038596423,0.0005040186,0.0003427758,0.0023258878,0.99538225,0.000024018373],"about_ca_topic_score_codex":0.0027872797,"about_ca_topic_score_gemma":0.008130524,"teacher_disagreement_score":0.052198905,"about_ca_system_score_codex":0.0024714638,"about_ca_system_score_gemma":0.0036921003,"threshold_uncertainty_score":0.17462271},"labels":[],"label_agreement":null},{"id":"W2145840143","doi":"10.1109/ase.2001.989790","title":"Wins and losses of algebraic transformations of software architectures","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Waterloo","funders":"","keywords":"Computer science; Abstraction; Graph; Software; Set (abstract data type); Architecture; Theoretical computer science; Programming language; Software architecture; Software engineering","score_opus":0.009522235452532233,"score_gpt":0.24226380770120534,"score_spread":0.23274157224867312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145840143","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5784863,0.0014657826,0.2829116,0.009566621,0.00031980756,0.00014649352,0.0001655878,0.0016976787,0.12523997],"genre_scores_gemma":[0.9624504,0.0005199707,0.024232876,0.00027483335,0.00015771654,0.00007077684,0.00008392223,0.00043510858,0.011774386],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99103105,0.0029070945,0.00039771872,0.0007982709,0.0037749945,0.0010907501],"domain_scores_gemma":[0.983334,0.008253983,0.0012429421,0.00525585,0.0011989179,0.00071429677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065604355,0.00062437746,0.0008456313,0.0015482479,0.0021692428,0.0052023204,0.0014940063,0.0016657027,0.005399605],"category_scores_gemma":[0.026074717,0.00064861553,0.0009903064,0.0009560225,0.0101726195,0.013636733,0.00735499,0.0046100407,0.0006845379],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025824388,0.000048674516,0.0011326942,0.00007005742,0.000027116159,0.00031600537,0.00079166476,0.010120222,0.00202763,0.94306123,0.0014815896,0.040664922],"study_design_scores_gemma":[0.00003248989,0.000069897505,0.00064922986,0.000020294327,0.00002530043,0.0003683816,0.0007854296,0.020362519,0.0038199027,0.9638804,0.009964166,0.00002194005],"about_ca_topic_score_codex":0.001127349,"about_ca_topic_score_gemma":0.0010352964,"teacher_disagreement_score":0.0065604355,"about_ca_system_score_codex":0.0018071722,"about_ca_system_score_gemma":0.0014143027,"threshold_uncertainty_score":0.034695268},"labels":[],"label_agreement":null},{"id":"W2145890718","doi":"10.1145/1985793.1985867","title":"Understanding broadcast based peer review on open source software projects","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":192,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Code review; Computer science; Process (computing); Interview; Open source software; Empirical research; Software peer review; Software development; Software; Software technical review; World Wide Web; Software engineering; Data science; Knowledge management; Software development process; Software quality; Software construction; Political science","score_opus":0.32245179144277475,"score_gpt":0.3380713667061482,"score_spread":0.01561957526337343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145890718","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7724114,0.0043416573,0.10568828,0.029715829,0.00043560631,0.000990359,0.00018086426,0.000338758,0.08589717],"genre_scores_gemma":[0.98316634,0.001038019,0.011004388,0.00072176213,0.00018382454,0.00026590028,0.00009956494,0.000107911525,0.0034123932],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.8567646,0.11011636,0.0049687354,0.0060098087,0.019731345,0.002409116],"domain_scores_gemma":[0.48655403,0.41643557,0.044227157,0.012607263,0.035175513,0.0050004437],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.07070141,0.0005749037,0.00072366314,0.006674054,0.0075155697,0.011745326,0.0031639661,0.0050469367,0.0026640391],"category_scores_gemma":[0.26694784,0.0011835148,0.0004885512,0.003286391,0.010573426,0.015481094,0.00834404,0.0035797907,0.0006883498],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000097380675,0.00009632209,0.017735006,0.00060455955,0.00004873096,0.0019419328,0.89850605,0.0007982622,0.0024114558,0.024682736,0.0035897947,0.04948776],"study_design_scores_gemma":[0.000092793794,0.00043398992,0.071592264,0.002068811,0.00012423092,0.0025576695,0.65216225,0.017474528,0.0025820893,0.06894432,0.18162932,0.00033772475],"about_ca_topic_score_codex":0.012686176,"about_ca_topic_score_gemma":0.01299429,"teacher_disagreement_score":0.98825467,"about_ca_system_score_codex":0.011431087,"about_ca_system_score_gemma":0.008350876,"threshold_uncertainty_score":0.3739093},"labels":[],"label_agreement":null},{"id":"W2146472972","doi":"10.1145/949344.949425","title":"*J","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Compiler; Java; Programming language; Metric (unit); Just-in-time compilation; Software engineering","score_opus":0.017030836973094623,"score_gpt":0.256896619143884,"score_spread":0.23986578217078938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146472972","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014563093,0.00086333824,0.39888173,0.0050777025,0.0043010144,0.00071489,0.011616071,0.033981703,0.53000045],"genre_scores_gemma":[0.10672127,0.0013885974,0.3823068,0.0030480307,0.00119778,0.00108412,0.02165525,0.010678696,0.4719195],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983668,0.00021461827,0.0001728194,0.0003797789,0.00068201317,0.00018390373],"domain_scores_gemma":[0.99674547,0.00033549615,0.00022085434,0.0011403991,0.0012387866,0.0003188563],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0012480519,0.00064776733,0.0005163372,0.0019411285,0.0020433248,0.003694276,0.0014250245,0.0009118245,0.12841702],"category_scores_gemma":[0.0059249494,0.00042103793,0.00063631433,0.0015729301,0.00071232487,0.004549283,0.0030704367,0.0013384797,0.070950076],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018706036,0.00017811044,0.004294063,0.00028644057,0.00003581574,0.00020604249,0.00068727345,0.00086731574,0.0047952435,0.16512066,0.38461784,0.43872416],"study_design_scores_gemma":[0.000015058327,0.000050166043,0.001716011,0.000057915422,0.000019311117,0.00032574212,0.00015627372,0.003139341,0.0036532376,0.032910492,0.9579114,0.000045068242],"about_ca_topic_score_codex":0.0032490904,"about_ca_topic_score_gemma":0.005656574,"teacher_disagreement_score":0.871583,"about_ca_system_score_codex":0.0006153596,"about_ca_system_score_gemma":0.0019178542,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2146503100","doi":"10.1109/msr.2007.35","title":"What Can OSS Mailing Lists Tell Us? A Preliminary Psychometric Text Analysis of the Apache Developer Mailing List","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":117,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"World Wide Web; Personality psychology; Computer science; Personality; Big Five personality traits; Baseline (sea); Open source; Open source software; Software; Information retrieval; Psychology; Operating system","score_opus":0.024407533881747698,"score_gpt":0.28822998575438064,"score_spread":0.2638224518726329,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146503100","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9912425,0.000110309324,0.0024443537,0.00063668244,0.000028000259,0.00013669793,0.0013754021,0.00009580386,0.0039301855],"genre_scores_gemma":[0.99167013,0.00013241642,0.004506052,0.00015964307,0.00006681877,0.00024352527,0.0021577654,0.00004416333,0.0010195231],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99628687,0.0014009521,0.0005802934,0.00040349996,0.0011127378,0.00021562254],"domain_scores_gemma":[0.8618884,0.09993789,0.01417973,0.00497673,0.01670619,0.002311032],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008598466,0.00038471658,0.00030568565,0.002584791,0.00064152275,0.002199236,0.0004513941,0.0007439322,0.0027478938],"category_scores_gemma":[0.08620424,0.00022390066,0.00045759152,0.002744586,0.00054551184,0.003785152,0.00073684135,0.0008175191,0.001249084],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048653944,0.0006465853,0.8779274,0.00029506005,0.0000684636,0.00018700611,0.013035536,0.00043549313,0.0038282431,0.00052722474,0.003078642,0.09948381],"study_design_scores_gemma":[0.000029285899,0.00059964816,0.9778553,0.00014153344,0.00005436566,0.00014317229,0.009678104,0.0031896816,0.001932907,0.0010982453,0.00522899,0.000048892452],"about_ca_topic_score_codex":0.0025605836,"about_ca_topic_score_gemma":0.003513266,"teacher_disagreement_score":0.008598466,"about_ca_system_score_codex":0.0005230069,"about_ca_system_score_gemma":0.0006046244,"threshold_uncertainty_score":0.045473576},"labels":[],"label_agreement":null},{"id":"W2146518621","doi":"10.1109/wcre.1995.514690","title":"Retrieving information from data flow diagrams","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Data flow diagram; Information retrieval; Semantic data model; Syntax; Component (thermodynamics); Diagram; Data mining; Natural language processing; Database","score_opus":0.05188803823137102,"score_gpt":0.24868717912186838,"score_spread":0.19679914089049738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146518621","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0097784,0.001499115,0.9523747,0.00090073526,0.00019247633,0.0004963717,0.008048833,0.014260484,0.012448889],"genre_scores_gemma":[0.056548357,0.0028403327,0.9105776,0.00023182396,0.00010746541,0.00033893154,0.02178355,0.0027997785,0.004772079],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971756,0.0005187281,0.00044450015,0.00037229506,0.0012851189,0.00020377597],"domain_scores_gemma":[0.98912454,0.005034819,0.0006137327,0.0023942038,0.0026113768,0.000221257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003001315,0.0021924148,0.0014349007,0.018929927,0.0017174557,0.005199562,0.0018355864,0.0022613283,0.008634147],"category_scores_gemma":[0.024454717,0.0010328257,0.0020290406,0.011656845,0.00091290334,0.010614727,0.003278145,0.002002502,0.005981699],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000272482,0.00022460833,0.0035425622,0.0026479657,0.000107289256,0.0018649715,0.003194471,0.008976391,0.02458332,0.13209623,0.04848596,0.77400374],"study_design_scores_gemma":[0.00012447312,0.00014581096,0.0022152446,0.0014033963,0.00030989866,0.0022618682,0.0017986622,0.06272619,0.079521164,0.24744944,0.60174006,0.00030382388],"about_ca_topic_score_codex":0.0054830797,"about_ca_topic_score_gemma":0.004532266,"teacher_disagreement_score":0.018929927,"about_ca_system_score_codex":0.0016185952,"about_ca_system_score_gemma":0.0033384648,"threshold_uncertainty_score":0.028884113},"labels":[],"label_agreement":null},{"id":"W2146648240","doi":"10.1109/icpc.2009.5090025","title":"Automatic classication of large changes into maintenance categories","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":119,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Waterloo","funders":"","keywords":"Commit; Computer science; Metadata; Categorization; Software maintenance; Task (project management); Programming language; Information retrieval; Software; Artificial intelligence; Database; Software system; World Wide Web; Engineering","score_opus":0.01189253959451926,"score_gpt":0.27347992245489783,"score_spread":0.2615873828603786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146648240","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7340001,0.0010588461,0.2454899,0.00048543155,0.0001985377,0.00042025183,0.002125129,0.011017169,0.0052046855],"genre_scores_gemma":[0.8709643,0.00019228915,0.12047243,0.000094774536,0.00007179128,0.00016036555,0.0050005536,0.00029884937,0.00274462],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99780136,0.0004170725,0.00023327139,0.000700354,0.00065889023,0.00018904536],"domain_scores_gemma":[0.98116904,0.009175386,0.0026004936,0.002175383,0.0044481284,0.00043155107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002146271,0.0008368041,0.0007425899,0.007089269,0.00070986047,0.0016417532,0.0012590477,0.0009679542,0.0012034584],"category_scores_gemma":[0.014813718,0.00025903867,0.00058396906,0.0021703152,0.00048201505,0.0024622153,0.0012281261,0.0014306756,0.0009297871],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005955386,0.0003201735,0.13593838,0.0003333117,0.00008396322,0.00029398515,0.001230695,0.007202826,0.01980274,0.0017656639,0.007945007,0.82448775],"study_design_scores_gemma":[0.00008001528,0.00054878514,0.1623039,0.00018707728,0.00016576878,0.0014871943,0.0018228979,0.74127567,0.061053842,0.0141968075,0.016723324,0.00015485263],"about_ca_topic_score_codex":0.0026974613,"about_ca_topic_score_gemma":0.004354744,"teacher_disagreement_score":0.007089269,"about_ca_system_score_codex":0.00067899434,"about_ca_system_score_gemma":0.0009185051,"threshold_uncertainty_score":0.011350691},"labels":[],"label_agreement":null},{"id":"W2146763372","doi":"10.1109/icsm.2000.883045","title":"C/C++ conditional compilation analysis using symbolic execution","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bell (Canada); Polytechnique Montréal","funders":"","keywords":"Computer science; Programming language; Header; Preprocessor; Compile time; Compiler; Source code; Task (project management); Symbolic execution; Software; Code (set theory); Directive; Software engineering; Set (abstract data type)","score_opus":0.023129237547132003,"score_gpt":0.28810123516612385,"score_spread":0.2649719976189918,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146763372","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08615528,0.00009649009,0.8955011,0.00013159384,0.000027608621,0.000096453565,0.0002888469,0.0122192,0.005483519],"genre_scores_gemma":[0.53217214,0.00010786051,0.46391097,0.000052869575,0.000020923797,0.0001293344,0.00076490355,0.0010145432,0.001826406],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989009,0.00023211098,0.00006933662,0.00018378344,0.000459861,0.00015403112],"domain_scores_gemma":[0.9954607,0.0028709057,0.0002976003,0.000647655,0.00064845907,0.00007467011],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007004748,0.00088137604,0.00052894076,0.0011040485,0.0006255371,0.00081237475,0.0009861859,0.0003799253,0.0048576156],"category_scores_gemma":[0.005309288,0.00025989313,0.0006284305,0.00092602515,0.0014799006,0.0010887303,0.0009143288,0.0008749077,0.0005963268],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087151484,0.00028874827,0.005788817,0.0010139175,0.00009686458,0.0010850994,0.0007552039,0.4472522,0.1090177,0.10179517,0.0076239468,0.32441074],"study_design_scores_gemma":[0.000038555427,0.000087999004,0.0009255138,0.000029528788,0.000029334808,0.00014983618,0.00006002751,0.9116411,0.06479681,0.018472534,0.003736124,0.0000326646],"about_ca_topic_score_codex":0.004831129,"about_ca_topic_score_gemma":0.0049721505,"teacher_disagreement_score":0.0048576156,"about_ca_system_score_codex":0.0007390027,"about_ca_system_score_gemma":0.0016833834,"threshold_uncertainty_score":0.016250372},"labels":[],"label_agreement":null},{"id":"W2147213196","doi":"10.1109/fie.2008.4720643","title":"Learning software engineering principles using open source software","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Software engineering; Computer science; Code refactoring; Code review; Software development; Programming language; Java; Source code; Software construction; Software; Coding (social sciences); Software evolution; Software quality","score_opus":0.05546599000304273,"score_gpt":0.27608576151580566,"score_spread":0.22061977151276294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147213196","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24677524,0.00047861692,0.6910378,0.0037113721,0.00014976719,0.00085568044,0.00028079847,0.0036490804,0.053061634],"genre_scores_gemma":[0.23738396,0.001955587,0.7313439,0.0008973614,0.000115694704,0.00069145893,0.000812383,0.0005439835,0.026255658],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984951,0.0003471266,0.000110845685,0.00025165148,0.00066587125,0.00012947609],"domain_scores_gemma":[0.99207485,0.0037958222,0.0006122317,0.0013288835,0.0012957492,0.00089248235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003423589,0.0010045666,0.00047029176,0.0015904398,0.0009801311,0.0033431498,0.0015140178,0.0009840422,0.0095157195],"category_scores_gemma":[0.011167321,0.0005612784,0.00066147797,0.0007301442,0.001055655,0.0047365194,0.003753056,0.0026735808,0.0035760403],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011011747,0.0037710739,0.017098043,0.00069663057,0.0000348205,0.00074689806,0.018032288,0.007439111,0.032485787,0.031173903,0.01773378,0.87067753],"study_design_scores_gemma":[0.00020046334,0.0028978637,0.04086301,0.0023199143,0.00008798472,0.0049827537,0.015435031,0.047940183,0.07827156,0.33035395,0.4762917,0.00035563234],"about_ca_topic_score_codex":0.00036821616,"about_ca_topic_score_gemma":0.0012576747,"teacher_disagreement_score":0.0095157195,"about_ca_system_score_codex":0.0007471168,"about_ca_system_score_gemma":0.0020687082,"threshold_uncertainty_score":0.03183329},"labels":[],"label_agreement":null},{"id":"W2147386665","doi":"10.1109/tse.2012.70","title":"A large-scale empirical study of just-in-time quality assurance","year":2012,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":699,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal","funders":"","keywords":"Computer science; Software quality assurance; Quality assurance; Source lines of code; Quality (philosophy); Software quality; Code review; Scale (ratio); Software; Empirical research; Software quality analyst; Software bug; Data science; Software engineering; Software development; Operations management; Engineering","score_opus":0.03357803100855729,"score_gpt":0.3186450643686469,"score_spread":0.28506703336008965,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147386665","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9964929,0.00035352763,0.0012871032,0.0003263917,0.000011803026,0.00003983928,0.0002604744,0.000023607097,0.0012042498],"genre_scores_gemma":[0.99857664,0.00013886501,0.00052927283,0.00008824062,0.0000129419195,0.000032100153,0.000356343,0.000013978826,0.00025147386],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9864699,0.0070611006,0.00082634226,0.0020013894,0.003177015,0.00046418345],"domain_scores_gemma":[0.5754482,0.33543867,0.046446085,0.0183386,0.018926969,0.0054015033],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018737236,0.0005504697,0.0005454523,0.0021346756,0.0007920336,0.0012109333,0.0017413202,0.0010712658,0.0023887372],"category_scores_gemma":[0.13169342,0.00044346103,0.00063359295,0.0028554518,0.0016246059,0.003665669,0.0011714441,0.0026083754,0.00078523334],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015156965,0.00089636527,0.9746695,0.0001372657,0.0002684458,0.00017970183,0.0021773754,0.0016357413,0.00023287194,0.00068156637,0.0023869884,0.016582625],"study_design_scores_gemma":[0.00003337729,0.00056070765,0.9794597,0.00009669846,0.00007268736,0.00028638434,0.0019236719,0.013597496,0.00038079298,0.00053733075,0.003017405,0.00003376645],"about_ca_topic_score_codex":0.0071114493,"about_ca_topic_score_gemma":0.006842744,"teacher_disagreement_score":0.018737236,"about_ca_system_score_codex":0.0011202066,"about_ca_system_score_gemma":0.0007029035,"threshold_uncertainty_score":0.0990932},"labels":[],"label_agreement":null},{"id":"W2147399912","doi":"","title":"A knowledge discovery methodology for the performance evaluation of scientific software","year":2000,"lang":"en","type":"article","venue":"Purdue e-Pubs (Purdue University System)","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Software mining; Computer science; Knowledge extraction; Data mining; Software; Domain knowledge; Process (computing); Domain (mathematical analysis); Machine learning; Artificial intelligence; Software system; Programming language; Software construction; Mathematics","score_opus":0.07466716047145118,"score_gpt":0.2937825527500882,"score_spread":0.21911539227863702,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147399912","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019506277,0.00038865348,0.9931485,0.00051642285,0.00003265275,0.00048645292,0.0006662785,0.0006940047,0.0021163777],"genre_scores_gemma":[0.0329131,0.00042359348,0.9635724,0.00020676512,0.000055569573,0.000767672,0.0012838878,0.000085913256,0.0006910768],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96764106,0.008602847,0.004031024,0.0042419373,0.014929298,0.0005538912],"domain_scores_gemma":[0.9487738,0.030946182,0.0033508572,0.008953045,0.0074339043,0.0005422438],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026388798,0.0018054378,0.0020210913,0.015111178,0.0015759621,0.0092694545,0.005348784,0.0023815597,0.0032150275],"category_scores_gemma":[0.063049264,0.0009669892,0.0039373217,0.010317021,0.0031057794,0.008331428,0.0038801983,0.0034000021,0.0018386679],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017668336,0.0006096312,0.006408865,0.0022549483,0.0006346133,0.000528074,0.0012340517,0.077736415,0.0042871297,0.3066745,0.008620893,0.590834],"study_design_scores_gemma":[0.00009861524,0.00033156143,0.0023813397,0.00089977024,0.00029819208,0.00061069237,0.00064498186,0.5329751,0.011901079,0.39350322,0.05618555,0.00016999597],"about_ca_topic_score_codex":0.004661637,"about_ca_topic_score_gemma":0.0031166524,"teacher_disagreement_score":0.9736112,"about_ca_system_score_codex":0.0038796454,"about_ca_system_score_gemma":0.0047924235,"threshold_uncertainty_score":0.13955897},"labels":[],"label_agreement":null},{"id":"W2147542998","doi":"10.1109/wcre.2003.1287250","title":"Completeness of a fact extractor","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Completeness (order theory); Extractor; Computer science; Programming language; Source code; Mathematics; Engineering","score_opus":0.03092212224234648,"score_gpt":0.2767919349140687,"score_spread":0.24586981267172225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147542998","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020465767,0.00061949785,0.96175224,0.0012410633,0.0001143072,0.0003499185,0.0031212978,0.007347345,0.0049886024],"genre_scores_gemma":[0.16011971,0.0011046443,0.8112729,0.0007093604,0.00021026586,0.00046543582,0.014219602,0.003739213,0.008158897],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9592485,0.008282726,0.0051298565,0.0049915086,0.02127813,0.001069347],"domain_scores_gemma":[0.86284655,0.051111944,0.0062188664,0.04724452,0.031819575,0.00075856224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026467042,0.0017381501,0.0018886171,0.007327227,0.0028088146,0.0062506446,0.0025683504,0.0021001236,0.0073117814],"category_scores_gemma":[0.086082995,0.0019876694,0.0032597696,0.0036792953,0.0031308837,0.015806815,0.0062101674,0.0035895633,0.0032472026],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086088805,0.00031006546,0.018751476,0.0026991852,0.00074044295,0.002044111,0.0048898,0.024348244,0.034143798,0.30055362,0.024890296,0.58576804],"study_design_scores_gemma":[0.0001542736,0.00036768493,0.006671002,0.001169225,0.0010805733,0.0031696274,0.0012389679,0.1217775,0.2130837,0.32579607,0.3250993,0.00039218733],"about_ca_topic_score_codex":0.0033057202,"about_ca_topic_score_gemma":0.0027228922,"teacher_disagreement_score":0.026467042,"about_ca_system_score_codex":0.0016000888,"about_ca_system_score_gemma":0.005998507,"threshold_uncertainty_score":0.13997275},"labels":[],"label_agreement":null},{"id":"W2147555693","doi":"10.1109/wcre.1995.514698","title":"Pattern matching for design concept localization","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Programming language; Code generation; Compiler; Source code; KPI-driven code analysis; Code (set theory); Redundant code; Software visualization; Matching (statistics); Reverse engineering; Fragment (logic); Static program analysis; Software development; Software; Theoretical computer science; Key (lock); Software construction","score_opus":0.04744163885835213,"score_gpt":0.2684536603072659,"score_spread":0.2210120214489138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147555693","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017794825,0.00034398664,0.9929054,0.00015426066,0.000042034593,0.00011134729,0.00016328934,0.0017651664,0.0027349994],"genre_scores_gemma":[0.04853597,0.00047335756,0.94395614,0.00016115773,0.000060676586,0.00037904174,0.00086916564,0.00037406155,0.005190449],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958788,0.0011361014,0.00035704518,0.0009041802,0.0014867565,0.00023703565],"domain_scores_gemma":[0.9953768,0.0020045703,0.00037416135,0.00161054,0.0005535387,0.00008040919],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030801438,0.0012296541,0.0013309511,0.0040592044,0.0010871679,0.002542355,0.00271295,0.00235528,0.01542461],"category_scores_gemma":[0.014549014,0.0009777355,0.0018036788,0.004663834,0.0019241698,0.0066741104,0.002848398,0.0018643102,0.0062852753],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020899692,0.000112327245,0.00085851585,0.0005551258,0.00009786589,0.00028121035,0.0003310309,0.026654057,0.007585894,0.26917574,0.011501378,0.6826378],"study_design_scores_gemma":[0.00008344059,0.00012551792,0.00059418,0.00021140963,0.00006916228,0.0008366757,0.00016670645,0.35067463,0.012823808,0.5644216,0.06993507,0.000057747136],"about_ca_topic_score_codex":0.0022565362,"about_ca_topic_score_gemma":0.0018026027,"teacher_disagreement_score":0.01542461,"about_ca_system_score_codex":0.0015025695,"about_ca_system_score_gemma":0.0016492184,"threshold_uncertainty_score":0.051600456},"labels":[],"label_agreement":null},{"id":"W2147753601","doi":"10.82308/22547","title":"Objective quantification of program behaviour using dynamic metrics","year":2004,"lang":"en","type":"article","venue":"eScholarship@McGill (McGill)","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fonds Québécois de la Recherche sur la Nature et les Technologies; Faculty of Graduate Studies and Research, University of Alberta; Natural Sciences and Engineering Research Council of Canada; McGill University","keywords":"Computer science; Concurrency; Metric (unit); Profiling (computer programming); Java; Data mining; Suite; TRACE (psycholinguistics); Programming language","score_opus":0.03554985863856908,"score_gpt":0.3002275510045447,"score_spread":0.2646776923659756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147753601","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3769531,0.0009889662,0.6047091,0.00031210578,0.00007045561,0.000914342,0.0023889886,0.0025441833,0.011118751],"genre_scores_gemma":[0.77662444,0.0005387392,0.21739112,0.00006710952,0.0000343942,0.0011550637,0.0021615361,0.0005221615,0.0015053367],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98689216,0.0038015181,0.0010262087,0.0016842087,0.006192795,0.00040304856],"domain_scores_gemma":[0.94745815,0.026407013,0.0090538915,0.008435983,0.008036777,0.000608258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009356667,0.0014228724,0.00069782557,0.004593297,0.00058337464,0.0026958198,0.0012119793,0.00067628134,0.0015417573],"category_scores_gemma":[0.046971396,0.00042691073,0.00050974503,0.0033453242,0.0016081447,0.003840606,0.0019901018,0.001150294,0.0003655508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086811726,0.0011153807,0.104252785,0.0025627222,0.0005589248,0.00025515127,0.0028455295,0.21217667,0.15321669,0.072226115,0.0035528182,0.44636914],"study_design_scores_gemma":[0.00009913664,0.004854593,0.13830072,0.0008747289,0.0003734907,0.00064411666,0.0028329003,0.5064772,0.23674664,0.07443145,0.033922132,0.00044295462],"about_ca_topic_score_codex":0.0010443849,"about_ca_topic_score_gemma":0.0011660546,"teacher_disagreement_score":0.009356667,"about_ca_system_score_codex":0.0012820642,"about_ca_system_score_gemma":0.0012693035,"threshold_uncertainty_score":0.04948336},"labels":[],"label_agreement":null},{"id":"W2147801848","doi":"","title":"Bug introducing changes: A case study with Android","year":2012,"lang":"","type":"article","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Android (operating system); Software bug; Computer science; Software maintenance; Software development; Security bug; Software engineering; Software; Data science; Computer security; Operating system; Software security assurance; Information security","score_opus":0.06350606863338436,"score_gpt":0.3367799795097603,"score_spread":0.2732739108763759,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147801848","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99519914,0.00029798865,0.0024780757,0.00048458268,0.000015971591,0.00013432797,0.00011895976,0.000049864764,0.0012209555],"genre_scores_gemma":[0.98929393,0.00046272753,0.008330478,0.00018848645,0.000021467758,0.00006256731,0.00011491866,0.000050950854,0.0014745713],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99716645,0.0013262352,0.00022990895,0.00031709738,0.0007327586,0.00022760456],"domain_scores_gemma":[0.971787,0.021793446,0.002686841,0.0013277971,0.0015702042,0.00083480007],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022731156,0.0006498394,0.0004464524,0.0023701254,0.002435582,0.0010340209,0.0011527307,0.0027507858,0.0007916349],"category_scores_gemma":[0.020884566,0.0005020205,0.0005897507,0.0016042641,0.0016831853,0.0015224331,0.0013836873,0.0014351542,0.00018513783],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006729702,0.004401385,0.34982848,0.0013497394,0.00026234705,0.23939249,0.22397067,0.0054963822,0.018388778,0.0032303403,0.005754131,0.1472523],"study_design_scores_gemma":[0.00034887812,0.004076926,0.5354896,0.0007951557,0.00057956105,0.18157287,0.17652142,0.024941567,0.025378304,0.0048511615,0.045040254,0.00040424668],"about_ca_topic_score_codex":0.010234348,"about_ca_topic_score_gemma":0.030469311,"teacher_disagreement_score":0.010234348,"about_ca_system_score_codex":0.0009954462,"about_ca_system_score_gemma":0.00074083486,"threshold_uncertainty_score":0.020349562},"labels":[],"label_agreement":null},{"id":"W2147981755","doi":"10.1109/wcre.2008.52","title":"A Hybrid Query Engine for the Structural Analysis of Java and AspectJ Programs","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; RDF query language; Query language; AspectJ; Sargable; Query optimization; Query expansion; Programming language; Query by Example; Web search query; Web query classification; Java; Program comprehension; Spatial query; Visualization; Information retrieval; Software; Data mining; Software system; Search engine","score_opus":0.022736736595450738,"score_gpt":0.26193482929112755,"score_spread":0.23919809269567682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147981755","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010904076,0.0001454926,0.95384,0.000118167605,0.000024486051,0.00017327968,0.0005790144,0.03291509,0.0013003378],"genre_scores_gemma":[0.1294983,0.00029801138,0.8534654,0.00025604235,0.00003646327,0.00036683373,0.0036161363,0.0064614415,0.006001348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979857,0.00025199997,0.00024570667,0.00036261493,0.0010226537,0.00013129762],"domain_scores_gemma":[0.9973748,0.0011775123,0.00016452868,0.00045652752,0.00069520774,0.00013144681],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028333177,0.0009140077,0.00096340146,0.0019499534,0.0004882828,0.0022536232,0.0019693926,0.0011023906,0.003660702],"category_scores_gemma":[0.0053944057,0.0007829932,0.0012080427,0.0014323941,0.0007831521,0.0036093483,0.0020822464,0.0010283921,0.0015408803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019344308,0.00043732105,0.008122257,0.0014416374,0.00037226334,0.0014667152,0.0032428554,0.018163256,0.27357134,0.08188729,0.030699203,0.57866144],"study_design_scores_gemma":[0.00043851626,0.0005543857,0.006843572,0.00022416408,0.00040440995,0.0027253071,0.00075810164,0.6302792,0.19257614,0.03436395,0.13050734,0.00032495064],"about_ca_topic_score_codex":0.004354554,"about_ca_topic_score_gemma":0.003961817,"teacher_disagreement_score":0.004354554,"about_ca_system_score_codex":0.0006176867,"about_ca_system_score_gemma":0.0011147822,"threshold_uncertainty_score":0.01498419},"labels":[],"label_agreement":null},{"id":"W2148234439","doi":"10.1109/scam.2007.19","title":"A Framework for Studying Clones In Large Software Systems","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"clone (Java method); Computer science; Cloning (programming); Linux kernel; Source code; Software system; Software maintenance; Software; Software engineering; Software development; Programming language; Operating system; Data mining; Biology","score_opus":0.03959469923058079,"score_gpt":0.3274427764674381,"score_spread":0.2878480772368573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148234439","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051430636,0.00052492943,0.99109405,0.00037718067,0.000015351818,0.00031065708,0.00055293617,0.0012594344,0.0007223483],"genre_scores_gemma":[0.028604858,0.00028049696,0.9693477,0.000051561943,0.000022499853,0.00047643206,0.00090287277,0.00007608024,0.00023751223],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9920155,0.0029359737,0.0010987349,0.0015976226,0.0020314932,0.00032065518],"domain_scores_gemma":[0.97213674,0.018575044,0.0030296291,0.0035660155,0.0019372418,0.00075525313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010595769,0.002258941,0.0018509219,0.01387437,0.0022259292,0.0055808565,0.00438752,0.0023346972,0.0022346387],"category_scores_gemma":[0.030972537,0.001432982,0.0044387835,0.012206327,0.0038168686,0.010903865,0.004222219,0.0027429357,0.0006365069],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025505526,0.00063725363,0.029996213,0.002735848,0.0007542667,0.0020059107,0.010232859,0.09086352,0.015239589,0.46052977,0.0104330275,0.37631667],"study_design_scores_gemma":[0.00011819297,0.00033127784,0.010279416,0.0005386357,0.00023392773,0.0018690228,0.0027697769,0.30460635,0.006427258,0.61265934,0.05993246,0.00023437433],"about_ca_topic_score_codex":0.008467145,"about_ca_topic_score_gemma":0.0068569845,"teacher_disagreement_score":0.01387437,"about_ca_system_score_codex":0.0022972734,"about_ca_system_score_gemma":0.003015692,"threshold_uncertainty_score":0.056036472},"labels":[],"label_agreement":null},{"id":"W2148607038","doi":"10.1109/icsm.1990.131315","title":"Composing subsystem structures using (k,2)-partite graphs","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Equivalence (formal languages); Computer science; Class (philosophy); Theoretical computer science; Software; Graph; Cluster analysis; Equivalence class (music); Programming language; Artificial intelligence; Discrete mathematics; Mathematics","score_opus":0.048419764036506725,"score_gpt":0.26802328477729126,"score_spread":0.21960352074078454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148607038","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014850201,0.000029276993,0.9773352,0.00004973255,0.000014300656,0.000073948795,0.000092668924,0.0015473976,0.006007329],"genre_scores_gemma":[0.16333662,0.00014248883,0.8233523,0.00008635727,0.000019222734,0.00018596121,0.0008100676,0.00091068266,0.011156306],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9991015,0.00024613997,0.000065030865,0.00018274573,0.00030750633,0.00009701268],"domain_scores_gemma":[0.99839693,0.00049374334,0.00013310045,0.0005728607,0.00031682145,0.000086534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000638452,0.0006483127,0.00043622524,0.001955626,0.0015977231,0.0015673895,0.0008734381,0.00062409724,0.005647071],"category_scores_gemma":[0.0032592784,0.0009742893,0.001186167,0.0016178773,0.00087410945,0.0027408742,0.001545581,0.00070064666,0.0021893824],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002193756,0.00017739572,0.003306131,0.0004615486,0.00020363345,0.0011183339,0.0030089861,0.21164218,0.05677428,0.30120018,0.010382944,0.41150492],"study_design_scores_gemma":[0.00004717312,0.00014092978,0.0017806813,0.00008343358,0.00017640063,0.00060289324,0.00070312386,0.61112213,0.04612648,0.27994913,0.059163716,0.000103844766],"about_ca_topic_score_codex":0.0053326734,"about_ca_topic_score_gemma":0.012349934,"teacher_disagreement_score":0.005647071,"about_ca_system_score_codex":0.000608071,"about_ca_system_score_gemma":0.00083759584,"threshold_uncertainty_score":0.018891335},"labels":[],"label_agreement":null},{"id":"W2148615889","doi":"10.1007/s10515-014-0162-2","title":"Automatic, high accuracy prediction of reopened bugs","year":2014,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Ministry of Science and Technology of the People's Republic of China; National Natural Science Foundation of China","keywords":"Software bug; Computer science; Eclipse; Rework; Software; Measure (data warehouse); Software engineering; Data mining; Operating system; Embedded system","score_opus":0.008804326722980867,"score_gpt":0.22699896113547427,"score_spread":0.2181946344124934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148615889","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57926786,0.0031880105,0.3487728,0.0008545072,0.000569269,0.00021085847,0.008069899,0.054132927,0.0049339244],"genre_scores_gemma":[0.85103554,0.00032510777,0.13597958,0.00017733125,0.00015277232,0.000054203854,0.007593747,0.0005828489,0.004098819],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99640226,0.00040860116,0.00024075697,0.00093883096,0.0017628314,0.0002466935],"domain_scores_gemma":[0.98318243,0.0072659007,0.0020978814,0.0028149227,0.0041237823,0.0005150625],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013159479,0.0012376868,0.0011252648,0.0031026436,0.0005039722,0.0015488532,0.0016108726,0.001726524,0.002136458],"category_scores_gemma":[0.01131249,0.00045181147,0.00072344113,0.0011644701,0.0002916335,0.0014016432,0.0012578329,0.0012492072,0.0027795725],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012518435,0.0009505948,0.122704886,0.0006141345,0.00031601978,0.0009794426,0.00020405115,0.050878495,0.08012531,0.0015240799,0.031996448,0.7084546],"study_design_scores_gemma":[0.00008504413,0.00028669563,0.029733034,0.000042347056,0.00010328478,0.0006845078,0.00006215627,0.91003203,0.05230376,0.0032298667,0.00337928,0.000058035224],"about_ca_topic_score_codex":0.003173891,"about_ca_topic_score_gemma":0.0056783613,"teacher_disagreement_score":0.003173891,"about_ca_system_score_codex":0.00041556745,"about_ca_system_score_gemma":0.0011190494,"threshold_uncertainty_score":0.007147193},"labels":[],"label_agreement":null},{"id":"W2148637904","doi":"10.1109/re.2005.44","title":"Modelling assumptions and requirements in the context of project risk","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Context (archaeology); Computer science; Cover (algebra); Risk analysis (engineering); Domain (mathematical analysis); Section (typography); Operations research; Engineering; Mathematics; Business","score_opus":0.07383924824698429,"score_gpt":0.32401875325933366,"score_spread":0.25017950501234937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148637904","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06066054,0.00066691096,0.91110265,0.003941112,0.0000817324,0.00031018796,0.00036431308,0.00028298906,0.022589535],"genre_scores_gemma":[0.69288486,0.0011195674,0.29977253,0.00023110435,0.00016572136,0.00080881105,0.00043018395,0.00010436626,0.0044828826],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9786825,0.013728358,0.0016897371,0.0016314693,0.003337623,0.00093024434],"domain_scores_gemma":[0.8918324,0.085649915,0.010978317,0.005429128,0.0048166383,0.0012936704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017635325,0.0015371039,0.0009427228,0.0034735892,0.0011961728,0.006343335,0.0029166278,0.004815727,0.0036943706],"category_scores_gemma":[0.075403705,0.001399679,0.0016468592,0.0030953533,0.004960388,0.010446525,0.0033886558,0.0037228595,0.00064697233],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007668459,0.00006616124,0.003587881,0.0002029895,0.000059522892,0.0009997665,0.0022615173,0.3131957,0.00039786365,0.66378033,0.000876711,0.014494955],"study_design_scores_gemma":[0.000045613324,0.00010517997,0.0010941817,0.00014874102,0.00004998081,0.00046526163,0.00070542743,0.44418135,0.000393191,0.5445365,0.008204935,0.00006970334],"about_ca_topic_score_codex":0.0052247136,"about_ca_topic_score_gemma":0.0028992293,"teacher_disagreement_score":0.017635325,"about_ca_system_score_codex":0.0031761427,"about_ca_system_score_gemma":0.002607124,"threshold_uncertainty_score":0.09326565},"labels":[],"label_agreement":null},{"id":"W2148787816","doi":"10.1145/2393596.2393661","title":"Seeking the ground truth","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":120,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Reuse; Set (abstract data type); Adaptation (eye); Java; Application programming interface; Software; Code (set theory); Recommender system; Programming language; Software engineering; Information retrieval; World Wide Web","score_opus":0.028865762426180653,"score_gpt":0.2707485777267796,"score_spread":0.24188281530059896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148787816","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24368352,0.01300312,0.5635697,0.07089046,0.0025803943,0.0005175407,0.014568468,0.0034221173,0.087764665],"genre_scores_gemma":[0.81610256,0.0027091608,0.1552004,0.004820576,0.00073946407,0.00031467437,0.014377615,0.00060471805,0.0051308763],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.982873,0.0060969265,0.0013187816,0.005616447,0.0033296347,0.0007651151],"domain_scores_gemma":[0.8891166,0.069600716,0.006869079,0.019037727,0.013754472,0.0016214706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012709544,0.0014697612,0.0015326986,0.0066464315,0.0026849716,0.011671527,0.003510893,0.005842259,0.02166692],"category_scores_gemma":[0.11938499,0.00090443256,0.0008727092,0.0050001936,0.004517184,0.019797994,0.004257517,0.004623157,0.008926254],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012147283,0.00067431026,0.06889697,0.003033741,0.00044034232,0.0011742517,0.0044709183,0.016708037,0.0051202807,0.34452137,0.08495405,0.468791],"study_design_scores_gemma":[0.00018392308,0.0003255822,0.012552109,0.002107011,0.0001995872,0.00096040574,0.008753952,0.10744081,0.007742449,0.7269227,0.13266253,0.00014893366],"about_ca_topic_score_codex":0.003335011,"about_ca_topic_score_gemma":0.0034407533,"teacher_disagreement_score":0.02166692,"about_ca_system_score_codex":0.002475285,"about_ca_system_score_gemma":0.003598484,"threshold_uncertainty_score":0.07248306},"labels":[],"label_agreement":null},{"id":"W2148854374","doi":"10.1145/2491411.2491444","title":"Convergent contemporary software peer review practices","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":301,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software peer review; Software review; Software; Android (operating system); Code review; World Wide Web; Software inspection; Software technical review; Data science; Software engineering; Software development; Software development process; Software quality; Software construction; Operating system","score_opus":0.07734132822184826,"score_gpt":0.33241864865250137,"score_spread":0.2550773204306531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148854374","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84913653,0.011026096,0.030144893,0.012486884,0.00026820626,0.00047534922,0.000100490295,0.00037671934,0.09598473],"genre_scores_gemma":[0.990333,0.0017083323,0.00417551,0.0005419049,0.00019725157,0.00011455568,0.000033373486,0.00007348223,0.0028226045],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8660275,0.054560505,0.006750944,0.016367212,0.052729405,0.0035644406],"domain_scores_gemma":[0.4985318,0.2698434,0.064395554,0.052390445,0.1017473,0.013091471],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06591981,0.00035089056,0.0007510504,0.0074772923,0.0062681525,0.01051146,0.003200315,0.0023306224,0.0035861828],"category_scores_gemma":[0.26128122,0.0006336101,0.0004807482,0.0050497414,0.011294031,0.006861081,0.0071992557,0.0023744544,0.0010158593],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035397272,0.00040218612,0.15129618,0.002250377,0.00038411192,0.001550717,0.3253777,0.002165177,0.008111614,0.081959106,0.010435963,0.41571283],"study_design_scores_gemma":[0.0002031581,0.0013278739,0.3875415,0.0028290395,0.00023867801,0.0042525264,0.17842795,0.008623498,0.0072462335,0.08211926,0.32667583,0.00051443017],"about_ca_topic_score_codex":0.0031978593,"about_ca_topic_score_gemma":0.0038924983,"teacher_disagreement_score":0.9340802,"about_ca_system_score_codex":0.007536173,"about_ca_system_score_gemma":0.008451839,"threshold_uncertainty_score":0.34862143},"labels":[],"label_agreement":null},{"id":"W2148965601","doi":"10.1109/wcre.2011.46","title":"Make it or Break it: Mining Anomalies from Linux Kbuild","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Linux kernel; Computer science; Kernel (algebra); Operating system; Source code; Consistency (knowledge bases); Configfs; System call; Anomaly detection; Code (set theory); Software bug; sysfs; Data mining; Programming language; Software; Artificial intelligence; Set (abstract data type)","score_opus":0.07296652176832846,"score_gpt":0.2807233845586774,"score_spread":0.20775686279034894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148965601","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94319475,0.00143676,0.033578403,0.00051490415,0.0001377518,0.00021103007,0.011039358,0.008145596,0.0017415644],"genre_scores_gemma":[0.9214653,0.00043398395,0.052435625,0.0001284859,0.00005346342,0.0001589015,0.023989167,0.0005264273,0.00080869294],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99557626,0.000337874,0.0005967484,0.0009995895,0.0020836908,0.00040587477],"domain_scores_gemma":[0.9842588,0.0066142417,0.0037862468,0.0019494382,0.0028772482,0.0005140718],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018586471,0.0009716155,0.0008225604,0.008965604,0.0009278888,0.0015963331,0.0015940162,0.0011204098,0.00034943878],"category_scores_gemma":[0.014391998,0.00045743474,0.0008200825,0.0060346317,0.00093929475,0.0020610883,0.0015367635,0.0010838194,0.00045277138],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005850184,0.00055311894,0.74242395,0.000931738,0.00030183457,0.0049580457,0.0037551923,0.017148087,0.014853611,0.0024683042,0.017569782,0.19445135],"study_design_scores_gemma":[0.000084988016,0.0003739637,0.5392723,0.00036467155,0.00044590587,0.006110344,0.006123887,0.35730633,0.041914776,0.011181027,0.036547996,0.0002738355],"about_ca_topic_score_codex":0.015971566,"about_ca_topic_score_gemma":0.016329922,"teacher_disagreement_score":0.015971566,"about_ca_system_score_codex":0.0011187178,"about_ca_system_score_gemma":0.0012837438,"threshold_uncertainty_score":0.031757236},"labels":[],"label_agreement":null},{"id":"W2149058489","doi":"10.1145/1066129.1066146","title":"Design mentoring based on design evolution analysis","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science","score_opus":0.03845726233902584,"score_gpt":0.27148832162381736,"score_spread":0.2330310592847915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149058489","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049650963,0.00024576831,0.91044945,0.0017785546,0.00006137292,0.00048986694,0.000055016113,0.001915482,0.035353452],"genre_scores_gemma":[0.4835461,0.0002796031,0.50508225,0.00029032285,0.000037612703,0.00047830094,0.00019879235,0.0002804384,0.00980641],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.989796,0.00631696,0.00035714792,0.0008712716,0.0022406322,0.00041792748],"domain_scores_gemma":[0.9635649,0.021842904,0.0027411405,0.0062557007,0.004670648,0.0009246362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010276121,0.0009485829,0.0004921685,0.0028218722,0.0015679724,0.002635929,0.0015869499,0.0014531084,0.0059472057],"category_scores_gemma":[0.04900019,0.00056571036,0.00080090587,0.001347662,0.001462018,0.0032001704,0.003278447,0.0018562397,0.00086569967],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018484207,0.0007013569,0.010577355,0.00030632244,0.00010662607,0.00049417024,0.0046463488,0.03935813,0.0064568263,0.11400758,0.0045821574,0.8185783],"study_design_scores_gemma":[0.00015432996,0.00065976084,0.0075438768,0.00040655295,0.0001592277,0.0010299801,0.0024713795,0.71449,0.018752884,0.19117674,0.06302475,0.00013044535],"about_ca_topic_score_codex":0.0013271378,"about_ca_topic_score_gemma":0.001921355,"teacher_disagreement_score":0.010276121,"about_ca_system_score_codex":0.0018735138,"about_ca_system_score_gemma":0.0023257444,"threshold_uncertainty_score":0.054345965},"labels":[],"label_agreement":null},{"id":"W2149155803","doi":"10.1109/csmr.2012.21","title":"Mining Kbuild to Detect Variability Anomalies in Linux","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Linux kernel; Lift (data mining); Code (set theory); Operating system; Kernel (algebra); Source code; Reliability (semiconductor); System call; Programming language; Data mining; Set (abstract data type)","score_opus":0.023205817912775234,"score_gpt":0.28162433149869703,"score_spread":0.2584185135859218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149155803","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9071693,0.0017297848,0.06540472,0.00044001217,0.00007961366,0.0001396789,0.013146558,0.0092368405,0.0026534675],"genre_scores_gemma":[0.9245805,0.0003007838,0.04878445,0.000055778026,0.00002554832,0.000108102635,0.024943061,0.00039951736,0.0008021796],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99796116,0.00023310458,0.00028317265,0.0005957322,0.00071068754,0.0002161114],"domain_scores_gemma":[0.9917956,0.0037492828,0.0016701375,0.0009837248,0.0015225784,0.00027869223],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012129147,0.00068992184,0.0007751129,0.0076488014,0.00081187364,0.001462117,0.001378679,0.00088822417,0.0006569802],"category_scores_gemma":[0.011778653,0.00046669054,0.00064658525,0.0047714124,0.0004764507,0.0017305649,0.0009293561,0.000644117,0.00055970694],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010950376,0.00069714396,0.5485315,0.0015734418,0.00033460595,0.0041615414,0.0026265753,0.043460026,0.029639468,0.004689677,0.020972053,0.34221905],"study_design_scores_gemma":[0.00015134642,0.00024329255,0.19754101,0.00022460107,0.00028476326,0.004689102,0.002183059,0.6898887,0.04548401,0.016694123,0.04242863,0.00018738028],"about_ca_topic_score_codex":0.010094017,"about_ca_topic_score_gemma":0.013321882,"teacher_disagreement_score":0.010094017,"about_ca_system_score_codex":0.0009501617,"about_ca_system_score_gemma":0.0013028819,"threshold_uncertainty_score":0.020070553},"labels":[],"label_agreement":null},{"id":"W2149200999","doi":"10.1109/icsm.2009.5306280","title":"Balancing value and modifiability when planning for the next release","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Perspective (graphical); Feature (linguistics); Software engineering; Domain (mathematical analysis); Software release life cycle; Resource (disambiguation); Software; Programming language; Software quality; Artificial intelligence; Software development; Mathematics","score_opus":0.038884468957516065,"score_gpt":0.2896279995859374,"score_spread":0.2507435306284213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149200999","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26240796,0.0009388509,0.71153337,0.0015767203,0.00010565115,0.0010458815,0.00015774647,0.00059592706,0.02163782],"genre_scores_gemma":[0.7646348,0.00027299623,0.23209949,0.000096521624,0.000042710675,0.00046347303,0.00013900624,0.00024785765,0.0020030716],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9840044,0.007954807,0.00091878173,0.0011688656,0.0051154196,0.00083763205],"domain_scores_gemma":[0.95369864,0.032724928,0.005243844,0.0024612702,0.004313933,0.0015573943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019528681,0.0013183326,0.0010517932,0.0032397227,0.001294521,0.0045038774,0.0016406056,0.0014552535,0.0024223502],"category_scores_gemma":[0.06605316,0.0010579837,0.0007414716,0.0019632492,0.0013440248,0.006962968,0.0020178237,0.0021987036,0.0003929525],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084315555,0.00049216603,0.023423962,0.00066657126,0.00029951727,0.00091483176,0.0022999712,0.30203938,0.023584064,0.047542684,0.0023942194,0.59549946],"study_design_scores_gemma":[0.0002074335,0.002760898,0.042212475,0.00041397996,0.0004110522,0.0010299553,0.002790148,0.7570925,0.026118698,0.15002917,0.01645084,0.00048281965],"about_ca_topic_score_codex":0.002999043,"about_ca_topic_score_gemma":0.0040719896,"teacher_disagreement_score":0.019528681,"about_ca_system_score_codex":0.0023904478,"about_ca_system_score_gemma":0.0029473521,"threshold_uncertainty_score":0.103278816},"labels":[],"label_agreement":null},{"id":"W2149423951","doi":"10.1109/icsme.2014.66","title":"Interactive Visualization of Bug Reports Using Topic Evolution and Extractive Summaries","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Software bug; Visualization; Software; Software engineering; Software maintenance; Security bug; Data science; Software visualization; Software development; Software construction; Data mining; Programming language; Software security assurance; Computer security","score_opus":0.01568546861017836,"score_gpt":0.29923977780621563,"score_spread":0.2835543091960373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149423951","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13338111,0.005372105,0.7055418,0.0013190376,0.00052708416,0.0007152461,0.014673403,0.13185424,0.0066159735],"genre_scores_gemma":[0.3785654,0.0030641877,0.5974889,0.00018101592,0.0003847542,0.0006300375,0.012147365,0.0035629147,0.003975418],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99898905,0.00024932803,0.0001545023,0.0002185543,0.00033223248,0.00005629863],"domain_scores_gemma":[0.99090344,0.0046814885,0.0012297252,0.0009363329,0.0019442264,0.00030489385],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001939579,0.0022586423,0.00094834605,0.007863571,0.0004297669,0.0024675229,0.0009295851,0.00087474816,0.0034271798],"category_scores_gemma":[0.009413808,0.00048474577,0.0008390796,0.0031927503,0.00019885522,0.0023439138,0.0013386168,0.00076702045,0.0012804955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018122214,0.0002534563,0.013277559,0.003282451,0.00040252638,0.000987634,0.006913585,0.008890512,0.08269003,0.0022540903,0.049530506,0.8297055],"study_design_scores_gemma":[0.00064520416,0.0015624965,0.0941056,0.0012283736,0.0013793712,0.0033370187,0.0062396177,0.4631172,0.17006914,0.011221518,0.24638043,0.00071401516],"about_ca_topic_score_codex":0.002332691,"about_ca_topic_score_gemma":0.0025932956,"teacher_disagreement_score":0.007863571,"about_ca_system_score_codex":0.00034604338,"about_ca_system_score_gemma":0.0005337048,"threshold_uncertainty_score":0.011465073},"labels":[],"label_agreement":null},{"id":"W2149598945","doi":"10.1109/wcre.2002.1173087","title":"On the use of metaballs to visually map source code structures and analysis results onto 3D space","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Source code; Computer science; Reverse engineering; KPI-driven code analysis; Software visualization; Visualization; Static program analysis; Code (set theory); Software; Point (geometry); Program comprehension; Focus (optics); Software engineering; Programming language; Software analytics; Space (punctuation); Human–computer interaction; Software system; Software development; Software construction; Data mining","score_opus":0.037226527322972325,"score_gpt":0.2853434246090273,"score_spread":0.24811689728605496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149598945","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031316234,0.00014470989,0.98543084,0.00036404008,0.000035180197,0.00004925558,0.00011311505,0.006508142,0.0042230096],"genre_scores_gemma":[0.03916341,0.00091830303,0.9535579,0.00019977755,0.00003178394,0.0001036981,0.0004264328,0.0021155209,0.0034831727],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989849,0.00034711557,0.00005944164,0.00013618829,0.0004266033,0.000045824167],"domain_scores_gemma":[0.99433315,0.0028470056,0.00028105153,0.0014369137,0.0009772304,0.00012465725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024245554,0.0014015295,0.0005112343,0.003283175,0.0009956368,0.0025776878,0.0017137994,0.0010043359,0.007979271],"category_scores_gemma":[0.010198219,0.0010485427,0.001191402,0.0020408332,0.0020918606,0.0037485498,0.004406089,0.0015197564,0.0030955386],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003504753,0.00009053742,0.0020467488,0.0005343428,0.00014164143,0.00087383547,0.0022802898,0.032403488,0.059059985,0.09630442,0.022025278,0.783889],"study_design_scores_gemma":[0.00015294338,0.00029403254,0.004006383,0.00068163965,0.00018448933,0.003486936,0.0007543927,0.46392632,0.105759345,0.15573832,0.26448765,0.00052748685],"about_ca_topic_score_codex":0.0056464635,"about_ca_topic_score_gemma":0.008593399,"teacher_disagreement_score":0.007979271,"about_ca_system_score_codex":0.00053365924,"about_ca_system_score_gemma":0.0008179714,"threshold_uncertainty_score":0.026693344},"labels":[],"label_agreement":null},{"id":"W2149672479","doi":"10.1145/1985793.1985842","title":"Non-essential changes in version histories","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":170,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Software evolution; Software engineering; Source code; Software; Software system; Code (set theory); Software development; Software construction; Programming language","score_opus":0.02151901268037202,"score_gpt":0.23720688623706596,"score_spread":0.21568787355669394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149672479","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9330776,0.0010687073,0.059093375,0.00016335525,0.00005607128,0.00014363218,0.0029590847,0.0007813783,0.0026568042],"genre_scores_gemma":[0.9804216,0.00021681473,0.016942393,0.000033261418,0.000032826414,0.00007146117,0.001662156,0.00010253487,0.00051688234],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9914595,0.0014688221,0.0012291797,0.0021344342,0.0033846134,0.00032342976],"domain_scores_gemma":[0.8510948,0.08865137,0.026660386,0.023154527,0.008840637,0.001598199],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006077871,0.00032100565,0.0005674896,0.007536908,0.00075687916,0.001331367,0.0008158653,0.0006181075,0.0010558129],"category_scores_gemma":[0.06731428,0.00054028316,0.00044988588,0.0066844714,0.00076609344,0.0021940803,0.0010710056,0.0009908354,0.0002593634],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042376647,0.00014191747,0.7936698,0.00047180793,0.00027883804,0.000651485,0.0031497553,0.0049501974,0.01542839,0.0053734654,0.0010832613,0.17437737],"study_design_scores_gemma":[0.00003285797,0.00027068588,0.91541433,0.00014314153,0.00017831761,0.0025696587,0.00078987796,0.038090974,0.018252587,0.011994179,0.012169579,0.000093757575],"about_ca_topic_score_codex":0.0015430731,"about_ca_topic_score_gemma":0.0033932503,"teacher_disagreement_score":0.007536908,"about_ca_system_score_codex":0.0006468124,"about_ca_system_score_gemma":0.0006642324,"threshold_uncertainty_score":0.032143235},"labels":[],"label_agreement":null},{"id":"W2149859327","doi":"10.1109/icsm.2009.5306306","title":"Refining clustering evaluation using structure indicators","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Computer science; Data mining; Software; Set (abstract data type); Decomposition; Machine learning","score_opus":0.038863057749679575,"score_gpt":0.33007175433872366,"score_spread":0.2912086965890441,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149859327","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36462682,0.0021552383,0.62148285,0.0006576438,0.00015154829,0.00069067796,0.0005849712,0.0021109062,0.0075393445],"genre_scores_gemma":[0.73027974,0.00047649967,0.26663816,0.00007449146,0.000054863947,0.0002875288,0.0010812794,0.0003276954,0.0007797672],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9645269,0.013900813,0.0027461417,0.0026568589,0.01501716,0.0011520558],"domain_scores_gemma":[0.86745316,0.07260053,0.012740839,0.008689998,0.036321435,0.0021940544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030669307,0.0020407944,0.0026227466,0.01888431,0.0017335246,0.0062844297,0.0018408685,0.002426778,0.0012318005],"category_scores_gemma":[0.14199898,0.0008420251,0.0014772979,0.008860707,0.002110239,0.006713825,0.0034319696,0.0017728455,0.0006424093],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021810557,0.00085715984,0.098573126,0.0017141406,0.00089683285,0.0002859968,0.0029183272,0.1950756,0.023939757,0.03143326,0.0054815444,0.6366432],"study_design_scores_gemma":[0.00022672428,0.0014851262,0.029272482,0.0004190845,0.00045303768,0.00039976626,0.0019621437,0.88926345,0.043931887,0.027539385,0.0047639366,0.00028302614],"about_ca_topic_score_codex":0.0029970848,"about_ca_topic_score_gemma":0.0030626801,"teacher_disagreement_score":0.030669307,"about_ca_system_score_codex":0.0032538923,"about_ca_system_score_gemma":0.0034465066,"threshold_uncertainty_score":0.1621967},"labels":[],"label_agreement":null},{"id":"W2149864547","doi":"10.1109/isese.2005.1541846","title":"Cloning by accident: an empirical study of source code cloning across software systems","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cloning (programming); Computer science; Source code; Domain (mathematical analysis); Set (abstract data type); Open source; Code (set theory); Programming language; Software; Notice; Operating system; World Wide Web; Software engineering","score_opus":0.031036959123870004,"score_gpt":0.35099501861555926,"score_spread":0.31995805949168926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149864547","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9951031,0.00027322603,0.0029104142,0.00028828272,0.000008403322,0.000104048755,0.000054660428,0.00002195126,0.0012358394],"genre_scores_gemma":[0.9968803,0.0002000597,0.0022905318,0.00014297759,0.000011474057,0.00009408183,0.00010604368,0.000027354317,0.0002471833],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.92236197,0.048185065,0.0066848206,0.0057884944,0.015054007,0.0019256534],"domain_scores_gemma":[0.40336218,0.45145702,0.08448038,0.024569945,0.031937446,0.0041930275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.044693876,0.0005048992,0.00075700076,0.0063155224,0.0042936965,0.0040783472,0.0029776052,0.0029875096,0.0015390349],"category_scores_gemma":[0.33561623,0.00092361216,0.00058819493,0.0066617145,0.008718396,0.01047858,0.0056444425,0.003634783,0.00037412823],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003646921,0.00071803853,0.74043894,0.0005670288,0.00016620295,0.0017761857,0.217898,0.00067909074,0.0009659247,0.003680837,0.0008609385,0.031884085],"study_design_scores_gemma":[0.00007536897,0.0022073376,0.63046277,0.00071395043,0.00022039794,0.0073292395,0.3220585,0.011431249,0.0033764294,0.0064441236,0.015449246,0.00023139284],"about_ca_topic_score_codex":0.006296062,"about_ca_topic_score_gemma":0.004856825,"teacher_disagreement_score":0.044693876,"about_ca_system_score_codex":0.0029008468,"about_ca_system_score_gemma":0.0024468773,"threshold_uncertainty_score":0.23636663},"labels":[],"label_agreement":null},{"id":"W2150113180","doi":"10.1109/icpc.2008.13","title":"Reading Beside the Lines: Indentation as a Proxy for Complexity Metric","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Metric (unit); Proxy (statistics); Computer science; Rank (graph theory); Indentation; Linguistic sequence complexity; Ranking (information retrieval); Code (set theory); Reading (process); Variance (accounting); Information retrieval; Mathematics; Programming language; Machine learning; Linguistics; Engineering","score_opus":0.07918617774518061,"score_gpt":0.33511092674346504,"score_spread":0.25592474899828443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150113180","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6386277,0.0012667065,0.33508834,0.00084819825,0.00029721507,0.0002983516,0.0047227154,0.00829291,0.010557937],"genre_scores_gemma":[0.8898151,0.00020230425,0.10488698,0.00009086621,0.0001450932,0.00015964833,0.002248878,0.0005602284,0.001891],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9893166,0.002368513,0.0012887211,0.0012144268,0.0053787483,0.00043310228],"domain_scores_gemma":[0.8337299,0.07959222,0.040345415,0.019964023,0.022615775,0.0037525715],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004939002,0.0008581656,0.0009626036,0.010902528,0.00066398864,0.0031145376,0.001337878,0.0011416189,0.0021893054],"category_scores_gemma":[0.09767747,0.0004955889,0.00062297075,0.010392075,0.001189727,0.005240477,0.0019220099,0.0014855355,0.0011461388],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009190995,0.00036562287,0.5454199,0.00090138527,0.00044439852,0.0007145923,0.0043564937,0.026385557,0.034095716,0.015172035,0.011341728,0.35988346],"study_design_scores_gemma":[0.000046549772,0.0012614898,0.6996285,0.00014235917,0.00018312536,0.0027486624,0.0020799353,0.1913343,0.051994644,0.026754929,0.023275869,0.00054956344],"about_ca_topic_score_codex":0.0017907966,"about_ca_topic_score_gemma":0.0031891966,"teacher_disagreement_score":0.010902528,"about_ca_system_score_codex":0.0010558818,"about_ca_system_score_gemma":0.0007764791,"threshold_uncertainty_score":0.026120245},"labels":[],"label_agreement":null},{"id":"W2150207557","doi":"10.1016/j.scico.2009.02.005","title":"Reading beside the lines: Using indentation to rank revisions by complexity","year":2009,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Indentation; Ranking (information retrieval); Code (set theory); Variance (accounting); Rank (graph theory); Metric (unit); Proxy (statistics); Simple (philosophy); Task (project management); Face (sociological concept); Reading (process); Algorithm; Theoretical computer science; Programming language; Artificial intelligence; Machine learning; Mathematics; Linguistics","score_opus":0.043479483258700706,"score_gpt":0.3439133840634598,"score_spread":0.30043390080475907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150207557","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30041504,0.0029930677,0.58898157,0.0071977116,0.0036818292,0.00080521073,0.0069094375,0.01118735,0.07782878],"genre_scores_gemma":[0.75334156,0.0005892205,0.22514533,0.00050331914,0.0014843342,0.0004040199,0.0049478943,0.0017389393,0.011845437],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9830283,0.007197151,0.0021902143,0.002005843,0.0044944845,0.001084014],"domain_scores_gemma":[0.77629775,0.12815645,0.023570897,0.0281378,0.037002645,0.006834509],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011495177,0.0017178081,0.0017778777,0.017667454,0.0032556944,0.008136672,0.0023745107,0.0020578338,0.03063122],"category_scores_gemma":[0.25678936,0.0006720479,0.0010590405,0.015443065,0.0037285571,0.01945807,0.005048565,0.0028766405,0.007332691],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018162047,0.00031060015,0.0457192,0.00067845907,0.00016072716,0.00039301228,0.0036837559,0.005901535,0.0025449838,0.0769192,0.08297163,0.7789006],"study_design_scores_gemma":[0.00048259064,0.0017758243,0.05469128,0.0010383893,0.00027170172,0.0013792831,0.008047144,0.12259056,0.013202911,0.6768135,0.11905979,0.00064705126],"about_ca_topic_score_codex":0.0020941412,"about_ca_topic_score_gemma":0.0027955223,"teacher_disagreement_score":0.03063122,"about_ca_system_score_codex":0.0012555897,"about_ca_system_score_gemma":0.0027247313,"threshold_uncertainty_score":0.10247165},"labels":[],"label_agreement":null},{"id":"W2150308295","doi":"10.1145/1453101.1453109","title":"Finding programming errors earlier by evaluating runtime monitors ahead-of-time","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":89,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; McGill University","funders":"","keywords":"Computer science; False positive paradox; Heap (data structure); Static analysis; Runtime verification; Debugging; Programming language; Benchmark (surveying); Alias; Set (abstract data type); Suite; Aliasing; Source code; Filter (signal processing); Formal verification; Data mining; Artificial intelligence","score_opus":0.04246659249662417,"score_gpt":0.31946929766524645,"score_spread":0.2770027051686223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150308295","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60293084,0.00069181976,0.36809528,0.0005891678,0.00016608599,0.00020351056,0.00080092123,0.023247236,0.0032752028],"genre_scores_gemma":[0.85175395,0.00013493156,0.14515999,0.00013776525,0.000033281565,0.00008128111,0.00080368726,0.0008213929,0.0010736386],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9940743,0.0012715553,0.00046601114,0.0014077553,0.002292194,0.0004883168],"domain_scores_gemma":[0.9728952,0.0119712595,0.0053639538,0.0054903454,0.0037633134,0.0005159573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004507312,0.0017584823,0.00082867115,0.0025623427,0.00045240382,0.0019089426,0.0016487953,0.0010751104,0.0011858857],"category_scores_gemma":[0.02979628,0.0007664275,0.00100844,0.0010395111,0.0008681861,0.0033589911,0.0011902307,0.0018134518,0.00049758755],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015252369,0.0008627254,0.26587525,0.00066696364,0.000546803,0.0010170534,0.0014973845,0.1858312,0.10968186,0.010343321,0.00531693,0.41683528],"study_design_scores_gemma":[0.000051963878,0.0006503267,0.021887837,0.000099228964,0.00024011821,0.00027842194,0.0002447157,0.8761037,0.08757091,0.008573309,0.0042129457,0.00008661725],"about_ca_topic_score_codex":0.0028774124,"about_ca_topic_score_gemma":0.004421631,"teacher_disagreement_score":0.004507312,"about_ca_system_score_codex":0.0010223024,"about_ca_system_score_gemma":0.0018061642,"threshold_uncertainty_score":0.023837268},"labels":[],"label_agreement":null},{"id":"W2150395559","doi":"10.1109/msr.2010.5463341","title":"The evolution of ANT build systems","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Agile software development; Executable; Software engineering; Software evolution; Software development; Software system; Source code; Software maintenance; Domain (mathematical analysis); Codebase; Software; Overhead (engineering); Perspective (graphical); Software construction; Programming language; Artificial intelligence","score_opus":0.0076390500866708394,"score_gpt":0.24046740806473182,"score_spread":0.23282835797806098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150395559","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9045707,0.0003626974,0.074146986,0.00064344855,0.00003911831,0.00012619635,0.00014234435,0.00073399843,0.019234471],"genre_scores_gemma":[0.9500127,0.00021099653,0.044185936,0.00008678808,0.000010164807,0.00007845414,0.00028944484,0.00016965192,0.004955897],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99859375,0.0004443903,0.00007568566,0.00023134139,0.00053162104,0.00012317703],"domain_scores_gemma":[0.9933241,0.0023158463,0.0010652202,0.0013310924,0.0014730421,0.0004906227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012774336,0.00025710152,0.00015354103,0.0009284308,0.000986871,0.0019577201,0.00074543525,0.00064268365,0.0010407853],"category_scores_gemma":[0.009761293,0.00035902436,0.0002496273,0.0006974167,0.0012791984,0.0014315718,0.0013546391,0.0007451476,0.00033580107],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005490641,0.00046433372,0.19441009,0.0006413463,0.00027329553,0.004011237,0.024569552,0.1616147,0.089194484,0.09290237,0.0076101106,0.42375943],"study_design_scores_gemma":[0.00009918691,0.0009280514,0.2109896,0.00020086297,0.00024002696,0.0052118828,0.009107391,0.4489089,0.029249322,0.05520224,0.23962262,0.00023982581],"about_ca_topic_score_codex":0.004064598,"about_ca_topic_score_gemma":0.0047308337,"teacher_disagreement_score":0.004064598,"about_ca_system_score_codex":0.0009875138,"about_ca_system_score_gemma":0.00085518043,"threshold_uncertainty_score":0.008081853},"labels":[],"label_agreement":null},{"id":"W2150439624","doi":"10.1002/smr.519","title":"Studying software evolution of large object‐oriented software systems using an ETGM algorithm","year":2010,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Computer science; Algorithm; Software; Scalability; Oracle; Software evolution; Software system; Theoretical computer science; Data mining; Software construction; Programming language","score_opus":0.017703981462783326,"score_gpt":0.29314088940744637,"score_spread":0.275436907944663,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150439624","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3401111,0.00017772726,0.6568596,0.00030564476,0.00001726245,0.00008440044,0.00011609255,0.0013475483,0.0009806176],"genre_scores_gemma":[0.35876977,0.00008134739,0.6398625,0.000042276955,0.000007505105,0.000067267596,0.0003241901,0.000093537994,0.00075167237],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99914014,0.0003030039,0.00006622858,0.00023199372,0.00021202222,0.00004661376],"domain_scores_gemma":[0.9936271,0.0042605954,0.0006168049,0.00061273895,0.0007411095,0.00014173286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026621886,0.00056914304,0.00046595445,0.0016422621,0.00045825826,0.00068569096,0.0010516625,0.0010721632,0.00077669474],"category_scores_gemma":[0.011850128,0.00031985765,0.000567743,0.001621885,0.00058648916,0.0015262854,0.00090890913,0.000695896,0.0001538146],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001404401,0.0001373414,0.018345049,0.00012044587,0.00007994186,0.0001593492,0.00033806838,0.81095386,0.008987635,0.0071808645,0.0007577187,0.15279934],"study_design_scores_gemma":[0.000009712981,0.00001640164,0.00081084867,0.0000037660363,0.0000070076494,0.000039335588,0.000019783798,0.9938471,0.0019269009,0.0030083458,0.0003079515,0.0000028192674],"about_ca_topic_score_codex":0.0056665298,"about_ca_topic_score_gemma":0.0038567425,"teacher_disagreement_score":0.0056665298,"about_ca_system_score_codex":0.0010601651,"about_ca_system_score_gemma":0.00085436186,"threshold_uncertainty_score":0.014079154},"labels":[],"label_agreement":null},{"id":"W2150526239","doi":"10.1109/ccece.2003.1226145","title":"A practical methodology for measurement deployment in GQM","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software deployment; Software; Systems engineering; Software engineering; Computer science; Software measurement; Engineering; Software development; Software construction","score_opus":0.3226216339184882,"score_gpt":0.424587630191256,"score_spread":0.10196599627276781,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150526239","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006885967,0.000028750597,0.99601775,0.00024310668,0.00005509246,0.00044002614,0.000027620821,0.0007738587,0.0017251971],"genre_scores_gemma":[0.011728732,0.00004461019,0.98630583,0.00006434358,0.00002146733,0.00089273916,0.00006491314,0.00016632365,0.00071105623],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.94246554,0.038771305,0.003693723,0.0042990497,0.010082666,0.0006876282],"domain_scores_gemma":[0.9599895,0.016953886,0.0021259373,0.011851753,0.008428601,0.0006503851],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.036184076,0.0017641953,0.00095054257,0.0046854503,0.002139064,0.0037164383,0.0030144206,0.002259215,0.008897522],"category_scores_gemma":[0.09125888,0.0014700545,0.0012210994,0.0044463277,0.0040178695,0.0063403933,0.006214769,0.0040899515,0.0040032812],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009518167,0.00029406027,0.0017257755,0.0011670775,0.00007848875,0.0003198387,0.007633794,0.0069938223,0.008826369,0.32669333,0.010557512,0.63561475],"study_design_scores_gemma":[0.00036852615,0.0011740997,0.0031880843,0.001967944,0.0001413447,0.0018536695,0.0047316784,0.12106632,0.024140917,0.44678533,0.394222,0.0003600902],"about_ca_topic_score_codex":0.0013648525,"about_ca_topic_score_gemma":0.0013382547,"teacher_disagreement_score":0.036184076,"about_ca_system_score_codex":0.0019755089,"about_ca_system_score_gemma":0.0033873825,"threshold_uncertainty_score":0.19136196},"labels":[],"label_agreement":null},{"id":"W2150628460","doi":"10.1109/csmr.2003.1192426","title":"A metric-based approach to enhance design quality through meta-pattern transformations","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":97,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Metric (unit); Task (project management); Transformation (genetics); Quality (philosophy); Object (grammar); Process (computing); Object-oriented programming; Object-oriented design; Metamodeling; Data mining; Software engineering; Artificial intelligence; Programming language; Systems engineering; Engineering","score_opus":0.14030917946349322,"score_gpt":0.3601159239258158,"score_spread":0.2198067444623226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150628460","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014686783,0.00029809924,0.97876054,0.00048769548,0.000042376138,0.00024710255,0.00007971453,0.0023984206,0.0029992769],"genre_scores_gemma":[0.14521345,0.00023015199,0.8527711,0.0000919054,0.000024270614,0.00023673786,0.0001612541,0.0002640226,0.0010071232],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99232304,0.0023383838,0.00066199194,0.0007188644,0.0037703824,0.00018731142],"domain_scores_gemma":[0.9861193,0.0032249743,0.002248478,0.0035453597,0.0045507713,0.0003112103],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056785815,0.0011824026,0.00094006606,0.0041591204,0.00062269944,0.0025149672,0.0021593897,0.00092880515,0.0011198279],"category_scores_gemma":[0.021532869,0.0005165371,0.000889734,0.0031013815,0.0013237686,0.0034269942,0.0013341035,0.001341139,0.00046173568],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023345376,0.00042086266,0.008533823,0.0006691098,0.00018647806,0.0002791263,0.0008308283,0.036307473,0.041105762,0.07205082,0.0036992426,0.83568305],"study_design_scores_gemma":[0.00018398154,0.0015681935,0.014702336,0.00034781758,0.00039151037,0.0023456304,0.0005229064,0.67118746,0.11317724,0.13451652,0.060769357,0.0002871714],"about_ca_topic_score_codex":0.0014436968,"about_ca_topic_score_gemma":0.0017417936,"teacher_disagreement_score":0.0056785815,"about_ca_system_score_codex":0.0013890358,"about_ca_system_score_gemma":0.0022556826,"threshold_uncertainty_score":0.030031562},"labels":[],"label_agreement":null},{"id":"W2150654203","doi":"10.1007/s10515-007-0007-3","title":"Differencing logical UML models","year":2007,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Correctness; Metamodeling; Computer science; Heuristics; Unified Modeling Language; Robustness (evolution); Programming language; Logical data model; Theoretical computer science; Data mining; Artificial intelligence; Software; Software engineering; Data modeling","score_opus":0.01960952468392616,"score_gpt":0.2525665858741463,"score_spread":0.23295706119022017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150654203","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049601756,0.0003139299,0.93730354,0.0003986317,0.00014686625,0.00013171059,0.00037368445,0.003337887,0.008392087],"genre_scores_gemma":[0.4696541,0.00031823033,0.5213138,0.00023444656,0.00003872327,0.00013596473,0.0015558513,0.00082702225,0.005921833],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9913816,0.002046888,0.00065419637,0.0010055942,0.004515569,0.0003960975],"domain_scores_gemma":[0.986201,0.005514468,0.00073220633,0.0053256033,0.0020727303,0.00015390867],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054809405,0.00069850957,0.00077449746,0.0029011667,0.0007025948,0.0043087616,0.0034365652,0.0016765792,0.006101684],"category_scores_gemma":[0.035241563,0.0011346365,0.0011819972,0.0023560992,0.0013876079,0.007250803,0.0055769123,0.0020859307,0.0017836909],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049878063,0.00022204188,0.005133696,0.00037401347,0.00011063031,0.0005571259,0.0014819467,0.060210712,0.030564258,0.5178201,0.003288559,0.37973815],"study_design_scores_gemma":[0.00008237917,0.00012558054,0.0012094349,0.0002024869,0.0001527424,0.00031578503,0.00059256324,0.5214901,0.06844986,0.3628802,0.044419326,0.00007956354],"about_ca_topic_score_codex":0.0023374131,"about_ca_topic_score_gemma":0.0038759653,"teacher_disagreement_score":0.006101684,"about_ca_system_score_codex":0.0021291978,"about_ca_system_score_gemma":0.0017723368,"threshold_uncertainty_score":0.028986335},"labels":[],"label_agreement":null},{"id":"W2150786161","doi":"10.1109/icsm.2005.91","title":"The top ten list: dynamic fault prediction","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":209,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Suite; Computer science; Heuristics; Quality (philosophy); Software; Focus (optics); Window (computing); Software quality; Software engineering; Database; Software development; Operating system","score_opus":0.006883087162793679,"score_gpt":0.2516377921719249,"score_spread":0.2447547050091312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150786161","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43210852,0.0019152263,0.51930994,0.0019936478,0.0003287772,0.00074160256,0.0066110706,0.027338527,0.009652597],"genre_scores_gemma":[0.6879267,0.00026364997,0.3032393,0.0002682319,0.00012191656,0.00019265799,0.0051224288,0.00036451247,0.00250059],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99855965,0.00025390106,0.000102782535,0.00028797498,0.000522276,0.00027356367],"domain_scores_gemma":[0.99025375,0.0047499365,0.0014751041,0.0008273693,0.0020315167,0.00066234596],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016845088,0.0013954293,0.0011173157,0.0049951957,0.0013124788,0.0020254962,0.0017799549,0.001226462,0.0028348768],"category_scores_gemma":[0.010315526,0.00050584215,0.0005257993,0.0028429409,0.00040082767,0.0025458047,0.0009864822,0.001064782,0.0010686172],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014052576,0.0010426946,0.10089899,0.00042109212,0.00023137669,0.000508281,0.00063416833,0.1887443,0.006705647,0.0050286986,0.055435512,0.638944],"study_design_scores_gemma":[0.00008296649,0.000299527,0.010182882,0.000050785704,0.000083837826,0.00027116144,0.0002778694,0.9699499,0.006493829,0.007280252,0.0049480484,0.000078972254],"about_ca_topic_score_codex":0.014854539,"about_ca_topic_score_gemma":0.02414095,"teacher_disagreement_score":0.014854539,"about_ca_system_score_codex":0.0011368662,"about_ca_system_score_gemma":0.0025891315,"threshold_uncertainty_score":0.029536188},"labels":[],"label_agreement":null},{"id":"W2150793216","doi":"10.1109/ictai.2008.149","title":"Mining Functional Aspects from Legacy Code","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Legacy system; Abstraction; Software engineering; Legacy code; Programming language; Software development; Code (set theory); Software; Object-oriented programming; Data science","score_opus":0.04617926157495618,"score_gpt":0.25342850410161893,"score_spread":0.20724924252666277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150793216","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6758793,0.0030224507,0.30582994,0.00051140256,0.00007547804,0.00043233868,0.0043697846,0.0034308939,0.0064483434],"genre_scores_gemma":[0.6466559,0.0018427623,0.3276676,0.00011152781,0.000097001335,0.00025738007,0.019143816,0.0006293979,0.0035946316],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99892575,0.00015265301,0.0001040221,0.00019877643,0.0005260683,0.00009257389],"domain_scores_gemma":[0.99172556,0.003235598,0.0013493724,0.0011780558,0.0022708713,0.000240571],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010543293,0.000993617,0.00044768056,0.007392216,0.0006369842,0.0010109209,0.00094776036,0.00081561547,0.00061500334],"category_scores_gemma":[0.008815716,0.0006216777,0.0010699348,0.0040433896,0.00047060274,0.0013920194,0.0010605681,0.00066324614,0.00038891492],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034769796,0.00029759307,0.15910804,0.0012171697,0.0002248024,0.005481742,0.0018215205,0.026104864,0.030179594,0.008928716,0.0068701063,0.7594181],"study_design_scores_gemma":[0.00012802478,0.00050702726,0.19771017,0.0005680803,0.0007261341,0.0074849636,0.0023239062,0.60741967,0.051154908,0.04919564,0.08256947,0.00021199246],"about_ca_topic_score_codex":0.004875376,"about_ca_topic_score_gemma":0.0084030265,"teacher_disagreement_score":0.007392216,"about_ca_system_score_codex":0.0003665686,"about_ca_system_score_gemma":0.0010912752,"threshold_uncertainty_score":0.00969398},"labels":[],"label_agreement":null},{"id":"W2150835538","doi":"10.1109/coginf.2011.6016167","title":"Validation of a generic GQM based measurement framework for software projects from industry practitioners","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Calgary; General Electric","keywords":"Software project management; Schedule; Computer science; Stakeholder; Project management; Software metric; Relevance (law); Metric (unit); Process management; Software development; Quality (philosophy); Project management triangle; Software; Software quality; Software engineering; Engineering management; Systems engineering; Engineering; Software construction; Operations management","score_opus":0.1425646335619076,"score_gpt":0.29897068390454556,"score_spread":0.15640605034263796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2150835538","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.574373,0.00020001484,0.40064183,0.0019684057,0.000101269245,0.008602043,0.00059178175,0.000542833,0.012978727],"genre_scores_gemma":[0.62852216,0.00011712219,0.3601415,0.0003092126,0.000014424717,0.008950321,0.0011260045,0.00009453932,0.0007246798],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9443333,0.037490834,0.0065627303,0.002521699,0.0077665546,0.0013249351],"domain_scores_gemma":[0.86021334,0.0712383,0.00716697,0.01697303,0.04252936,0.0018789062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08743232,0.00086047954,0.0006506389,0.0045526354,0.0016936687,0.0025710086,0.0018495067,0.0018536646,0.0012396348],"category_scores_gemma":[0.17191324,0.00043303127,0.0007712021,0.0037199392,0.0018238281,0.003406157,0.0035390279,0.0016005582,0.0003924555],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060665037,0.0048196185,0.20146614,0.0031948145,0.00017337136,0.00063110905,0.10730914,0.021302817,0.029885627,0.066198744,0.006782726,0.5576292],"study_design_scores_gemma":[0.000910381,0.009501873,0.49725676,0.004288125,0.00037000442,0.001583417,0.10260563,0.19100669,0.027582303,0.07406053,0.090308204,0.0005261234],"about_ca_topic_score_codex":0.0037144031,"about_ca_topic_score_gemma":0.0037249897,"teacher_disagreement_score":0.08743232,"about_ca_system_score_codex":0.004128326,"about_ca_system_score_gemma":0.008637442,"threshold_uncertainty_score":0.4623918},"labels":[],"label_agreement":null},{"id":"W2151181659","doi":"10.5555/2337223.2337318","title":"Temporal analysis of API usage concepts","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Application programming interface; Documentation; Software; Reuse; Software documentation; Set (abstract data type); Software engineering; Software development; Usage data; Programming language; Database; Software development process; World Wide Web","score_opus":0.028364928366075633,"score_gpt":0.32227386587443396,"score_spread":0.29390893750835834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151181659","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.784974,0.0010573333,0.20432067,0.00029403818,0.00003683167,0.00025570538,0.0022315362,0.0009429829,0.0058868383],"genre_scores_gemma":[0.92648065,0.00020975793,0.07098883,0.000032273063,0.000020339772,0.00018586997,0.0012931972,0.00007988558,0.00070911687],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.996725,0.0006686013,0.0003188902,0.0006620501,0.0014205475,0.00020479219],"domain_scores_gemma":[0.9827708,0.009565599,0.0025441574,0.0013829244,0.0033998531,0.0003366879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020632227,0.00039752942,0.00037625086,0.006919057,0.0006153317,0.001137053,0.0008468154,0.00041791334,0.0014381141],"category_scores_gemma":[0.015555033,0.00030016358,0.0007056944,0.0059899637,0.0006211222,0.0018981466,0.00091786345,0.0008015469,0.0002602574],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009773636,0.0006035445,0.31903166,0.0010736425,0.00035939857,0.0019559076,0.006280184,0.019053595,0.054882493,0.0342102,0.003907805,0.5576642],"study_design_scores_gemma":[0.00007967266,0.00062958046,0.2385125,0.0003548686,0.00031046098,0.0054332973,0.004215918,0.62637687,0.045847576,0.049637813,0.028433725,0.0001677125],"about_ca_topic_score_codex":0.0055723633,"about_ca_topic_score_gemma":0.0048491494,"teacher_disagreement_score":0.006919057,"about_ca_system_score_codex":0.00079059374,"about_ca_system_score_gemma":0.00085760077,"threshold_uncertainty_score":0.011079848},"labels":[],"label_agreement":null},{"id":"W2151277370","doi":"10.1109/icsm.2006.61","title":"Source-Level Linkage: Adding Semantic Information to C++ Fact-bases","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Executable; Compiler; Programming language; Source code; Semantics (computer science); Software; Template; Task (project management); Graph; Software engineering; Theoretical computer science; Engineering; Systems engineering","score_opus":0.018728781476592737,"score_gpt":0.24654764412065083,"score_spread":0.2278188626440581,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151277370","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031353552,0.0002569326,0.93141854,0.00050936063,0.000073535106,0.00038212538,0.005347964,0.024906648,0.0057513006],"genre_scores_gemma":[0.12596416,0.00028999284,0.8563104,0.00026827143,0.00005562544,0.0002473054,0.011426468,0.0035781583,0.0018596189],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9978053,0.00042656105,0.0002740934,0.00041710556,0.0009399377,0.00013701138],"domain_scores_gemma":[0.9842328,0.00830253,0.0015756915,0.0037142546,0.0019865944,0.000188105],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003274571,0.000955843,0.0006128612,0.0072596357,0.0009865757,0.0033899688,0.0019288118,0.0011225487,0.004689419],"category_scores_gemma":[0.022334177,0.0010585806,0.001336861,0.005158759,0.0014389552,0.0064086923,0.0027083945,0.0018630683,0.0014675424],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054713857,0.0003490216,0.021468375,0.0014201673,0.00024328938,0.0015680889,0.0035830454,0.060384113,0.01763609,0.13989425,0.027120288,0.7257862],"study_design_scores_gemma":[0.00015638483,0.00026424156,0.009379361,0.00063071615,0.0003307234,0.001444247,0.0010862146,0.41147953,0.094137296,0.21094918,0.26984134,0.00030066594],"about_ca_topic_score_codex":0.0044691227,"about_ca_topic_score_gemma":0.006120303,"teacher_disagreement_score":0.0072596357,"about_ca_system_score_codex":0.0012042508,"about_ca_system_score_gemma":0.0025141144,"threshold_uncertainty_score":0.017317832},"labels":[],"label_agreement":null},{"id":"W2151298976","doi":"10.1145/643603.643622","title":"Navigating and querying code without getting lost","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":235,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Code (set theory); Task (project management); Process (computing); Representation (politics); Software engineering; Source code; Human–computer interaction; Programming language; Systems engineering; Engineering","score_opus":0.01975542242781014,"score_gpt":0.29878952918632523,"score_spread":0.2790341067585151,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151298976","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13840212,0.00048924924,0.80963343,0.0010210054,0.00008382242,0.00012424638,0.00050309574,0.040290613,0.009452348],"genre_scores_gemma":[0.4280964,0.00052670017,0.5593664,0.00031726336,0.000037169266,0.000121761106,0.0010615201,0.004090932,0.006381956],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99871135,0.00030512243,0.00009403521,0.00022460247,0.0005823995,0.00008244417],"domain_scores_gemma":[0.9919703,0.0049219825,0.00032538074,0.0016156369,0.0008205769,0.00034608535],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001963923,0.0008286099,0.00079054816,0.0019530859,0.0008404852,0.0026418879,0.0017541577,0.0019748511,0.0042359736],"category_scores_gemma":[0.016256912,0.00050155324,0.00051440817,0.0012132198,0.0008931434,0.0056899446,0.0036134028,0.0012976998,0.0014480664],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007067541,0.00030830535,0.01122561,0.0007056351,0.00012916094,0.0024428824,0.011398214,0.009694561,0.07992092,0.027178897,0.022804484,0.83348465],"study_design_scores_gemma":[0.00028769748,0.0007183097,0.009307545,0.000620849,0.00030494903,0.0060366523,0.0074821934,0.48438844,0.15916859,0.10847547,0.2227812,0.00042811112],"about_ca_topic_score_codex":0.0021701036,"about_ca_topic_score_gemma":0.0038152244,"teacher_disagreement_score":0.0042359736,"about_ca_system_score_codex":0.00039939012,"about_ca_system_score_gemma":0.00076000695,"threshold_uncertainty_score":0.014170766},"labels":[],"label_agreement":null},{"id":"W2151580048","doi":"10.1145/1117696.1117710","title":"ConcernMapper","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":89,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; McGill University","keywords":"Computer science; Eclipse; Extensibility; Software engineering; Variety (cybernetics); Programming language; Plug-in; Interface (matter); Architecture; Separation of concerns; User interface; Human–computer interaction; World Wide Web; Operating system; Software; Artificial intelligence","score_opus":0.01736274973286346,"score_gpt":0.26901276067090346,"score_spread":0.25165001093804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151580048","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050085266,0.000762832,0.7025016,0.00075007696,0.00057249435,0.0004286956,0.003410514,0.25719404,0.029371219],"genre_scores_gemma":[0.088262446,0.0019680294,0.70696497,0.0031036683,0.00028301767,0.0011508999,0.023815183,0.09553984,0.07891196],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99847597,0.00020904544,0.00012005043,0.0002863286,0.00075736264,0.00015112535],"domain_scores_gemma":[0.9973369,0.0007410231,0.00014533101,0.0009291171,0.0006349396,0.00021271587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019021428,0.0016542556,0.00064190314,0.0010958585,0.0006113012,0.0026945202,0.0028985692,0.0014362835,0.023212072],"category_scores_gemma":[0.0070069167,0.0014309979,0.0014533147,0.0006407095,0.00059203414,0.005693296,0.0037265758,0.0029227175,0.015875721],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009212432,0.00034843982,0.00467519,0.0017855574,0.0001804036,0.0009187909,0.00096035853,0.0061773923,0.025810987,0.08377345,0.43287793,0.4415702],"study_design_scores_gemma":[0.00007881655,0.000084790205,0.00085414905,0.000100703546,0.00005579509,0.0008879614,0.00006571805,0.0160295,0.018306023,0.01618428,0.94727916,0.000073049836],"about_ca_topic_score_codex":0.0024590502,"about_ca_topic_score_gemma":0.003533948,"teacher_disagreement_score":0.023212072,"about_ca_system_score_codex":0.0007207577,"about_ca_system_score_gemma":0.001449972,"threshold_uncertainty_score":0.07765216},"labels":[],"label_agreement":null},{"id":"W2151666881","doi":"10.5555/2337223.2337510","title":"Online sharing and integration of results from mining software repositories","year":2012,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Cloud computing; Software; World Wide Web; Data science; Domain (mathematical analysis); Software as a service; Software mining; Source code; Software engineering; Data mining; Database; Software development; Information retrieval; Software construction","score_opus":0.04690364729029656,"score_gpt":0.29605498929087065,"score_spread":0.2491513420005741,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151666881","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50217974,0.005032356,0.2671525,0.0062690023,0.0007312355,0.0019232209,0.12989457,0.031488314,0.05532904],"genre_scores_gemma":[0.5116946,0.0018054994,0.23793656,0.00051219517,0.00042874677,0.00076685514,0.23759359,0.002349872,0.006912074],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98121256,0.0040826285,0.0022171678,0.0030676736,0.008595245,0.00082468527],"domain_scores_gemma":[0.909366,0.024491934,0.006623637,0.044996276,0.011778786,0.002743258],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.0140503235,0.001225777,0.0015967725,0.026449595,0.0018270733,0.007004373,0.0025624158,0.0015076236,0.0027530014],"category_scores_gemma":[0.06189361,0.00092668406,0.0015299945,0.020757778,0.0007466787,0.012664454,0.010818168,0.001733334,0.00245862],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014594403,0.0015820228,0.14568251,0.0021760897,0.0019480344,0.001476224,0.0069903946,0.010076412,0.017506735,0.017417172,0.069541946,0.7241431],"study_design_scores_gemma":[0.00030216447,0.00082862057,0.27972865,0.001468955,0.0015711054,0.0024686041,0.010585139,0.13968033,0.077457316,0.10264516,0.3824144,0.0008495284],"about_ca_topic_score_codex":0.002717406,"about_ca_topic_score_gemma":0.0039206683,"teacher_disagreement_score":0.9974376,"about_ca_system_score_codex":0.0011634035,"about_ca_system_score_gemma":0.0025027809,"threshold_uncertainty_score":0.07430607},"labels":[],"label_agreement":null},{"id":"W2151781530","doi":"10.1109/csmr.2009.62","title":"Software Clustering Using Dynamic Analysis and Static Dependencies","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Cluster analysis; Computer science; Program comprehension; Static analysis; Data mining; Software; Software system; Decomposition; Static program analysis; Software development; Artificial intelligence; Programming language","score_opus":0.016626119411495006,"score_gpt":0.28423898177960893,"score_spread":0.26761286236811394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151781530","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028982194,0.00018930824,0.96643305,0.00011710042,0.000014738103,0.00009057255,0.00010063323,0.001330516,0.0027418588],"genre_scores_gemma":[0.34969154,0.00035154732,0.64513546,0.000047283716,0.00004250078,0.00021465724,0.000939424,0.0005172894,0.0030603593],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985013,0.00028613224,0.000086102445,0.0003769569,0.0006327915,0.000116794065],"domain_scores_gemma":[0.9955388,0.0013906735,0.0007506658,0.0007266053,0.0014544684,0.00013881578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010705267,0.0009342494,0.0007343755,0.008688314,0.0015424414,0.0017519713,0.0011216342,0.00077306293,0.0015726892],"category_scores_gemma":[0.0064741788,0.0007372733,0.0011990286,0.004508429,0.0010015621,0.0024980952,0.0014808071,0.00072022073,0.0006893924],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015026754,0.00017750393,0.016378796,0.00034618488,0.0002217014,0.0004227775,0.0014382872,0.1872969,0.038869508,0.057325307,0.0042352425,0.6931376],"study_design_scores_gemma":[0.000018334642,0.00008354156,0.007694771,0.000068514266,0.00012711831,0.00051765906,0.00038877878,0.9152621,0.015360944,0.050191958,0.010195598,0.000090662834],"about_ca_topic_score_codex":0.005191998,"about_ca_topic_score_gemma":0.0072542494,"teacher_disagreement_score":0.008688314,"about_ca_system_score_codex":0.0014872454,"about_ca_system_score_gemma":0.0016246453,"threshold_uncertainty_score":0.010790706},"labels":[],"label_agreement":null},{"id":"W2151833352","doi":"10.1109/wcre.1997.624578","title":"The Orphan Adoption problem in architecture maintenance","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Toronto","funders":"","keywords":"Computer science; Architecture; Key (lock); Cluster analysis; Field (mathematics); Software engineering; Artificial intelligence; Computer security; Mathematics","score_opus":0.012469097838939049,"score_gpt":0.21969227304978314,"score_spread":0.2072231752108441,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151833352","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17770374,0.0014811744,0.81191593,0.0024133963,0.00016436224,0.00020507773,0.00010375592,0.0012114752,0.004801051],"genre_scores_gemma":[0.5634138,0.0012986971,0.4283105,0.0004828736,0.00017488906,0.00023538913,0.000248049,0.00033459146,0.005501237],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99332994,0.0020923896,0.00051638624,0.0014453173,0.0019456152,0.0006702946],"domain_scores_gemma":[0.979035,0.011861098,0.0035246818,0.0025534306,0.0023684658,0.00065727724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00555083,0.00093224423,0.0017073195,0.0025294914,0.0026392294,0.0022601709,0.0021133851,0.0033474842,0.0032090642],"category_scores_gemma":[0.037524264,0.0011006078,0.0014418194,0.0031998816,0.0028083287,0.008495331,0.0049941763,0.0035672078,0.0005184675],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078088464,0.0005075615,0.042069834,0.001027089,0.00044147106,0.00652321,0.0051579885,0.10122706,0.01619165,0.21278106,0.013008226,0.6002839],"study_design_scores_gemma":[0.00014827409,0.0007954435,0.0064304965,0.00037711064,0.00045404574,0.009607851,0.0020211176,0.60391045,0.012587164,0.336589,0.026933871,0.00014522191],"about_ca_topic_score_codex":0.0017789156,"about_ca_topic_score_gemma":0.0026425086,"teacher_disagreement_score":0.00555083,"about_ca_system_score_codex":0.0011243096,"about_ca_system_score_gemma":0.0013186139,"threshold_uncertainty_score":0.029355943},"labels":[],"label_agreement":null},{"id":"W2151877107","doi":"10.1109/tse.2004.69","title":"A cognitive-based mechanism for constructing software inspection teams","year":2004,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software inspection; Mechanism (biology); Process (computing); Selection (genetic algorithm); Software; Cognition; Artificial intelligence; Software engineering; Software development; Software quality","score_opus":0.015596019644818788,"score_gpt":0.24884721502036167,"score_spread":0.23325119537554287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151877107","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039371423,0.000091650785,0.93879735,0.0008729801,0.00008930394,0.00042158418,0.000037819005,0.0011351957,0.01918266],"genre_scores_gemma":[0.48908058,0.00007565448,0.50294495,0.00023079752,0.00007980851,0.00082699314,0.00009308821,0.00007363953,0.006594443],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935916,0.0021915745,0.00041006625,0.0013222591,0.001941071,0.0005434114],"domain_scores_gemma":[0.98645455,0.004596057,0.0023674546,0.0027009936,0.0024229556,0.0014580963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008816103,0.0009072998,0.0004892036,0.0037473408,0.0030147084,0.0045898077,0.0043542194,0.002488966,0.0070862086],"category_scores_gemma":[0.029339818,0.00070987706,0.0013852997,0.0015624013,0.0037532786,0.0047761244,0.00468692,0.0015773808,0.0014689539],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026974603,0.00070541847,0.009203213,0.00024583362,0.00023653563,0.00055333617,0.006779901,0.03747665,0.016881244,0.7016523,0.0051294523,0.22086638],"study_design_scores_gemma":[0.0003881864,0.0009899476,0.0058379094,0.00018770681,0.0002458149,0.00092889677,0.0020912339,0.34709403,0.013832598,0.5870764,0.041011747,0.00031542347],"about_ca_topic_score_codex":0.00231458,"about_ca_topic_score_gemma":0.0021352211,"teacher_disagreement_score":0.008816103,"about_ca_system_score_codex":0.0021903715,"about_ca_system_score_gemma":0.0034584042,"threshold_uncertainty_score":0.0466246},"labels":[],"label_agreement":null},{"id":"W2151949169","doi":"10.1109/iwpse.2005.8","title":"Change Impact Analysis for Requirement Evolution using Use Case Maps","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Change impact analysis; Computer science; Software engineering; Slicing; Software requirements specification; Requirements analysis; Process (computing); Software system; Software; Dependency (UML); Software evolution; Systems analysis; Systems engineering; Software construction; Engineering; Programming language","score_opus":0.11847484409562624,"score_gpt":0.34942318833603697,"score_spread":0.23094834424041072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151949169","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17359214,0.00020106057,0.8152291,0.0002783013,0.000019168992,0.0010066091,0.0005016417,0.0015541923,0.0076177404],"genre_scores_gemma":[0.6295708,0.00014944767,0.36758512,0.00003607767,0.000015143076,0.0009871104,0.0008233863,0.00013080848,0.00070208404],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9873329,0.0055186497,0.0006627557,0.0005453612,0.0055737607,0.00036655567],"domain_scores_gemma":[0.9401631,0.04857517,0.0028098302,0.0031155262,0.005030139,0.0003062746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008390869,0.0011667346,0.00071944384,0.010415767,0.0011409259,0.0019299532,0.0012326257,0.0011999587,0.002920423],"category_scores_gemma":[0.047595438,0.0007505582,0.00197488,0.0047158506,0.0010787931,0.002661581,0.0014731818,0.0011403097,0.00031017847],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007974112,0.00086813205,0.070709795,0.0010705452,0.000736361,0.0035488878,0.0047581256,0.30182067,0.024898842,0.06473251,0.0026807706,0.52337795],"study_design_scores_gemma":[0.000045553064,0.00023601911,0.01944992,0.00009373528,0.00019593448,0.0007985338,0.00086131215,0.9445661,0.011147388,0.016678799,0.0058330144,0.000093795374],"about_ca_topic_score_codex":0.007709781,"about_ca_topic_score_gemma":0.0051491363,"teacher_disagreement_score":0.010415767,"about_ca_system_score_codex":0.0019278436,"about_ca_system_score_gemma":0.0016601803,"threshold_uncertainty_score":0.044375718},"labels":[],"label_agreement":null},{"id":"W2152061679","doi":"10.1145/1321211.1321246","title":"Removing manually generated boilerplate from electronic texts","year":2007,"lang":"en","type":"article","venue":"Proceedings of CASCON","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; University of New Brunswick","funders":"","keywords":"Computer science; Boilerplate text; Metadata; Parsing; Template; Information retrieval; ASCII; World Wide Web; Natural language processing; Programming language","score_opus":0.008669633716138785,"score_gpt":0.244718544950474,"score_spread":0.2360489112343352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2152061679","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12884833,0.0015900052,0.7484623,0.0017124179,0.0027911195,0.0022344962,0.029623246,0.06893142,0.015806621],"genre_scores_gemma":[0.11721917,0.00062901684,0.7892835,0.00059079955,0.0004087764,0.0013406747,0.058973603,0.013706771,0.017847631],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9905324,0.0021050395,0.0013584982,0.0024167784,0.0032339387,0.00035326902],"domain_scores_gemma":[0.9283932,0.031704772,0.002947759,0.02168307,0.014741881,0.0005293427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060078143,0.0020441846,0.0020309929,0.009767527,0.0018788922,0.0036122939,0.0020775197,0.0017184722,0.010010262],"category_scores_gemma":[0.049325604,0.0018734997,0.0015910389,0.0090822475,0.0018524827,0.0032515547,0.0037148849,0.002720064,0.017494077],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000707271,0.00034822675,0.011949075,0.0029524239,0.0001815809,0.0028706607,0.0046394183,0.0049469997,0.06409403,0.0096029,0.1059959,0.79171145],"study_design_scores_gemma":[0.00022148847,0.0003113653,0.030597063,0.0008575229,0.00042849392,0.004376694,0.0032771158,0.087084256,0.28050828,0.021311568,0.5706049,0.00042123202],"about_ca_topic_score_codex":0.004627495,"about_ca_topic_score_gemma":0.0065574218,"teacher_disagreement_score":0.010010262,"about_ca_system_score_codex":0.0014211641,"about_ca_system_score_gemma":0.0035965608,"threshold_uncertainty_score":0.033487678},"labels":[],"label_agreement":null},{"id":"W2152570911","doi":"10.5121/ijsea.2010.1301","title":"Bayesian Network Based XP Process Modelling","year":2010,"lang":"en","type":"article","venue":"International Journal of Software Engineering & Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Regina","funders":"","keywords":"Extreme programming; Bayesian network; Computer science; Process (computing); Software; Software development; Machine learning; Software development process; Programming language","score_opus":0.008758295651533717,"score_gpt":0.2609209078791803,"score_spread":0.2521626122276466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2152570911","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013703592,0.00032670342,0.97657055,0.00038173975,0.00004433471,0.0001361148,0.0006433138,0.00038109918,0.007812563],"genre_scores_gemma":[0.7398399,0.0021157665,0.22676133,0.00022352286,0.00011992876,0.0014641411,0.0021172946,0.00016501406,0.027193025],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979615,0.00078367605,0.000115264076,0.0003958617,0.00057709304,0.0001665889],"domain_scores_gemma":[0.99599445,0.0027520943,0.00042896927,0.0001188852,0.00063301576,0.00007263828],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025862232,0.0010828782,0.0013681464,0.0021949532,0.00064920366,0.002423029,0.0027330923,0.0028645287,0.0060565947],"category_scores_gemma":[0.009380765,0.0010424262,0.0012545903,0.0026613378,0.0009025069,0.0026266673,0.0010463165,0.001974381,0.00131832],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031455784,0.000025564546,0.00046888634,0.000038219438,0.000024742822,0.00007232041,0.000045657296,0.9703621,0.0002460931,0.020806557,0.0003775996,0.0075007332],"study_design_scores_gemma":[0.0000067446094,0.000006623328,0.00010477009,0.0000057699162,0.000007801643,0.000013156043,0.0000035705582,0.9937748,0.0000627845,0.0055967383,0.00041123378,0.0000059133235],"about_ca_topic_score_codex":0.023658048,"about_ca_topic_score_gemma":0.012523249,"teacher_disagreement_score":0.023658048,"about_ca_system_score_codex":0.0023009782,"about_ca_system_score_gemma":0.0014848794,"threshold_uncertainty_score":0.04704064},"labels":[],"label_agreement":null},{"id":"W2152720480","doi":"10.1109/se.2007.9","title":"A Requirement Level Modification Analysis Support Framework","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Source code; Identification (biology); Software engineering; Software maintenance; Code (set theory); Risk analysis (engineering); Decision support system; Software; Software system; Data mining; Programming language; Set (abstract data type)","score_opus":0.09072435874838497,"score_gpt":0.3575111483323302,"score_spread":0.26678678958394525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2152720480","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025352803,0.00011805208,0.98554325,0.00070703257,0.00002802066,0.0003064974,0.00023224881,0.005283318,0.005246319],"genre_scores_gemma":[0.07251514,0.00017899254,0.92293674,0.00024385413,0.0000650263,0.0004069677,0.00066820905,0.00036742474,0.0026176781],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99139637,0.0019774144,0.0007849555,0.00087987725,0.004361438,0.0005999389],"domain_scores_gemma":[0.986999,0.0052388352,0.0013656234,0.0023814137,0.0035804491,0.00043465052],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010603339,0.0017378184,0.0009761457,0.004945927,0.0011892414,0.005135782,0.0047390107,0.0026416464,0.0049436986],"category_scores_gemma":[0.017604966,0.0012955181,0.0021057562,0.0018528044,0.0015788856,0.0050645517,0.002628188,0.0027833064,0.0026096995],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024400195,0.00048095314,0.003483822,0.0009351272,0.00023145172,0.002449796,0.0022255098,0.07449814,0.021947846,0.48053327,0.018662421,0.3943076],"study_design_scores_gemma":[0.00010365688,0.00021576542,0.0011205431,0.000493153,0.00017797058,0.0018650667,0.00044087437,0.6157437,0.023481032,0.17395057,0.18220182,0.00020587341],"about_ca_topic_score_codex":0.0048120013,"about_ca_topic_score_gemma":0.0034410558,"teacher_disagreement_score":0.010603339,"about_ca_system_score_codex":0.0017314597,"about_ca_system_score_gemma":0.0043383883,"threshold_uncertainty_score":0.056076467},"labels":[],"label_agreement":null},{"id":"W2152828763","doi":"10.1145/1342211.1342218","title":"Tracking code clones in evolving software","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"clone (Java method); Code refactoring; Computer science; Software maintenance; Software; Code (set theory); Cloning (programming); Tracking (education); Source code; Software evolution; Software system; Software development; Programming language; Software engineering; Data mining; Software construction; Biology; Genetics; Gene","score_opus":0.038347238878211405,"score_gpt":0.27418674510735275,"score_spread":0.23583950622914135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2152828763","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37621114,0.00033944566,0.6149461,0.00013303282,0.000030710537,0.00033453538,0.00018498792,0.006664207,0.001155866],"genre_scores_gemma":[0.48007563,0.00024607155,0.51734865,0.0000753407,0.00001107266,0.00013980777,0.00047472445,0.00038657908,0.0012421198],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.996497,0.0007093014,0.0003065436,0.00092218036,0.001447872,0.00011711719],"domain_scores_gemma":[0.9693066,0.01485096,0.0048331814,0.006340163,0.0042177816,0.00045137655],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030894761,0.0004586829,0.00064038683,0.0023300902,0.0007767669,0.00154613,0.0013568385,0.0012671008,0.0005002133],"category_scores_gemma":[0.02660832,0.00053422747,0.0004253683,0.0017800786,0.0008444223,0.0021969352,0.0018636849,0.0010105218,0.0002464501],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042220336,0.0003378397,0.048899002,0.00041852164,0.00009185717,0.0009593999,0.0043277307,0.019676145,0.17383868,0.008806068,0.0014829027,0.7407397],"study_design_scores_gemma":[0.00011811829,0.0011916028,0.040541735,0.00019520658,0.00031336746,0.0033731214,0.0010420597,0.5221532,0.3935552,0.015179818,0.02211565,0.00022102322],"about_ca_topic_score_codex":0.0022517487,"about_ca_topic_score_gemma":0.0025176995,"teacher_disagreement_score":0.0030894761,"about_ca_system_score_codex":0.0006284341,"about_ca_system_score_gemma":0.0009807284,"threshold_uncertainty_score":0.016338885},"labels":[],"label_agreement":null},{"id":"W2152976736","doi":"10.1109/icpc.2006.6","title":"A Metric-Based Heuristic Framework to Detect Object-Oriented Design Flaws","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Heuristics; Object-oriented design; Metric (unit); Object-oriented programming; Engineering design process; Heuristic; Source code; Structural pattern; Software engineering; Software design; Artificial intelligence; Programming language; Software development; Software; Engineering","score_opus":0.016487262181982442,"score_gpt":0.26465441789729444,"score_spread":0.248167155715312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2152976736","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021449137,0.00026436042,0.97573847,0.00014619404,0.000016008165,0.00022217032,0.0001235186,0.0009963522,0.0010438809],"genre_scores_gemma":[0.24819474,0.00012375717,0.75042,0.00007598153,0.0000209717,0.0003244679,0.00038866716,0.00008173572,0.0003696261],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9927671,0.0021686037,0.000742958,0.0008524642,0.0031825383,0.0002862752],"domain_scores_gemma":[0.98217815,0.008112743,0.0033206742,0.0021285152,0.0038389734,0.00042094613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069403276,0.0014209679,0.0013843476,0.008498488,0.00081197656,0.0026071093,0.0022773014,0.0012506334,0.0008026337],"category_scores_gemma":[0.028071666,0.0005181911,0.0010443059,0.0029597722,0.0016511903,0.002037341,0.001301592,0.0010060032,0.00029529195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004046675,0.0006942298,0.049030147,0.0011450232,0.000460103,0.00050872937,0.000977175,0.21908744,0.025558183,0.061013814,0.0037031523,0.6374173],"study_design_scores_gemma":[0.00008036859,0.00044783996,0.007193522,0.00012599361,0.00015548643,0.0005410173,0.0002409916,0.93033814,0.011864977,0.044350117,0.004548057,0.0001134709],"about_ca_topic_score_codex":0.0031981496,"about_ca_topic_score_gemma":0.0038078916,"teacher_disagreement_score":0.008498488,"about_ca_system_score_codex":0.002057018,"about_ca_system_score_gemma":0.0026445917,"threshold_uncertainty_score":0.03670436},"labels":[],"label_agreement":null},{"id":"W2153256103","doi":"10.1016/j.scico.2014.05.008","title":"A survey of grammatical inference in software engineering","year":2014,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Grammar induction; Inference; Rule-based machine translation; Artificial intelligence; Grammar; Natural language processing; Variety (cybernetics); Software; Programming language; Class (philosophy); Linguistics","score_opus":0.022185499959028836,"score_gpt":0.28292738078839835,"score_spread":0.2607418808293695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153256103","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017321875,0.3316728,0.57513446,0.014899513,0.0011160073,0.00016125533,0.0011620698,0.0014169111,0.05711519],"genre_scores_gemma":[0.19378139,0.42213434,0.3580458,0.005531684,0.005525473,0.0002708794,0.0036736412,0.0013770501,0.009659723],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9945332,0.0022089165,0.0007108037,0.0010505329,0.0012960428,0.00020057286],"domain_scores_gemma":[0.96716523,0.027643204,0.000608031,0.0020711701,0.002291136,0.00022116785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067884983,0.0008625108,0.0019302784,0.008422466,0.0017196073,0.0053037032,0.0024986167,0.002607226,0.007988327],"category_scores_gemma":[0.029710878,0.0011664262,0.0012725585,0.011550422,0.0035391708,0.017308,0.002590143,0.0028498864,0.00280845],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007699445,0.00011810449,0.0038716816,0.0034566121,0.00009535296,0.00018862248,0.0008644193,0.0037205326,0.0012181666,0.27264404,0.014572364,0.69917315],"study_design_scores_gemma":[0.000029090465,0.00009294755,0.0038954085,0.0022457577,0.00012992375,0.0011461035,0.00082192913,0.02637982,0.0031313202,0.65378577,0.30824262,0.00009932126],"about_ca_topic_score_codex":0.0024692828,"about_ca_topic_score_gemma":0.002629411,"teacher_disagreement_score":0.008422466,"about_ca_system_score_codex":0.002044167,"about_ca_system_score_gemma":0.0032190168,"threshold_uncertainty_score":0.035901427},"labels":[],"label_agreement":null},{"id":"W2153258502","doi":"10.1109/icre.2003.1232745","title":"Improving requirements tracing via information retrieval","year":2004,"lang":"en","type":"article","venue":"Journal of Lightwave Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":290,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University; National Aeronautics and Space Administration","keywords":"Tracing; Traceability; Computer science; Framing (construction); Focus (optics); Information retrieval; Data mining; Software engineering; Programming language; Engineering","score_opus":0.011368775472954637,"score_gpt":0.25395738827077696,"score_spread":0.24258861279782232,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153258502","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014866353,0.00062897033,0.9719909,0.00050944794,0.000046426012,0.00026675008,0.00015504248,0.009953854,0.0015822479],"genre_scores_gemma":[0.07424155,0.0005007153,0.92180663,0.00025703304,0.00009037374,0.00019860307,0.0007945309,0.00044655596,0.0016640644],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98384064,0.0062129153,0.001445903,0.0019522193,0.0059165694,0.0006318164],"domain_scores_gemma":[0.9302547,0.037933018,0.005400044,0.012692853,0.0132194515,0.00049991556],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012247896,0.0029562975,0.0028301012,0.015327729,0.0019320457,0.00532977,0.0038652576,0.003144339,0.0027867695],"category_scores_gemma":[0.072859295,0.0011694885,0.0023368218,0.009241696,0.0012383431,0.00910394,0.0034065812,0.0029681807,0.0029461095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021262617,0.0005344162,0.002949537,0.0006128333,0.00015148016,0.00016689234,0.00086754566,0.016602727,0.023289109,0.004731224,0.0065511744,0.9433305],"study_design_scores_gemma":[0.00029166654,0.0009082921,0.005742236,0.000375075,0.0006880193,0.0017244188,0.0010617081,0.7615347,0.15753168,0.03559711,0.034200005,0.0003450458],"about_ca_topic_score_codex":0.0074626985,"about_ca_topic_score_gemma":0.006462133,"teacher_disagreement_score":0.015327729,"about_ca_system_score_codex":0.0015269513,"about_ca_system_score_gemma":0.004117292,"threshold_uncertainty_score":0.0647738},"labels":[],"label_agreement":null},{"id":"W2153344509","doi":"10.1145/1810295.1810337","title":"Bridging lightweight and heavyweight task organization","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Task management; Task (project management); Bridging (networking); Categorization; Software development; Software; Software engineering; Task analysis; Work (physics); Human–computer interaction; Knowledge management; Systems engineering; Artificial intelligence; Engineering","score_opus":0.004728723637470368,"score_gpt":0.21323280508432765,"score_spread":0.20850408144685728,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153344509","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2199209,0.00054760947,0.76971734,0.0010199357,0.00019101112,0.0004487407,0.000080266545,0.0016421784,0.0064319912],"genre_scores_gemma":[0.6107483,0.00026980543,0.3799293,0.00048327292,0.00011349947,0.0005960369,0.00026966573,0.0011400483,0.0064499346],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97643924,0.010912226,0.0024477392,0.0032566956,0.00455408,0.0023900366],"domain_scores_gemma":[0.8825776,0.0546898,0.012677835,0.0397935,0.0060940804,0.0041671824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020087702,0.001154748,0.0011934782,0.0029483836,0.0032277515,0.0065901806,0.0034961514,0.002450136,0.0027587556],"category_scores_gemma":[0.08302361,0.0013929105,0.0008792694,0.0026083172,0.005329368,0.01280417,0.021797875,0.0029601762,0.0018628264],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019154795,0.00080464815,0.041916285,0.0013207932,0.00013861309,0.0017822265,0.05629267,0.008798392,0.05051812,0.17329195,0.00527393,0.6579469],"study_design_scores_gemma":[0.00028717567,0.0014953081,0.030718578,0.0011760903,0.00026547455,0.0029005152,0.024003461,0.11695106,0.03892143,0.6537392,0.12904038,0.00050138764],"about_ca_topic_score_codex":0.0016954553,"about_ca_topic_score_gemma":0.0021152243,"teacher_disagreement_score":0.020087702,"about_ca_system_score_codex":0.0019384699,"about_ca_system_score_gemma":0.0035870178,"threshold_uncertainty_score":0.106235206},"labels":[],"label_agreement":null},{"id":"W2153538963","doi":"10.1109/wpc.2000.852478","title":"A pattern matching framework for software architecture recovery and restructuring","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Pattern matching; Restructuring; Software architecture; Matching (statistics); Data mining; Software; Software system; Architecture; Abstraction; Theoretical computer science; Artificial intelligence; Programming language; Mathematics","score_opus":0.018186454397399802,"score_gpt":0.24671636242118156,"score_spread":0.22852990802378176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153538963","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00031791578,0.0002136041,0.99729556,0.0002193513,0.000029986302,0.00009160385,0.00009315143,0.0006208109,0.0011179659],"genre_scores_gemma":[0.008041759,0.0003527378,0.9893454,0.00010261713,0.000044962064,0.00022035932,0.00038264887,0.000095241536,0.0014143254],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9950197,0.0013023865,0.00065826287,0.0009424811,0.0018135918,0.00026346577],"domain_scores_gemma":[0.9976792,0.000759098,0.00029893982,0.00066604506,0.00046150337,0.0001352233],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050977184,0.0017954472,0.0014690832,0.0053066197,0.0020438868,0.00434885,0.004125359,0.0026008638,0.0059544016],"category_scores_gemma":[0.007994768,0.0010341204,0.003622627,0.0059449878,0.0031980705,0.00737626,0.003093779,0.0030886184,0.0023617651],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005652134,0.00007970772,0.00042069083,0.00043807653,0.00008907452,0.00044650768,0.00039780705,0.027703518,0.0028366372,0.7432214,0.01063349,0.21367648],"study_design_scores_gemma":[0.000045378343,0.00008936357,0.0002464327,0.00024790427,0.00007459033,0.0009469045,0.00017653282,0.15672989,0.003331505,0.728765,0.10928748,0.000058894435],"about_ca_topic_score_codex":0.0076874983,"about_ca_topic_score_gemma":0.006435434,"teacher_disagreement_score":0.0076874983,"about_ca_system_score_codex":0.002090763,"about_ca_system_score_gemma":0.003764589,"threshold_uncertainty_score":0.026959598},"labels":[],"label_agreement":null},{"id":"W2153546999","doi":"10.1145/1189748.1189751","title":"Representing concerns in source code","year":2007,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":215,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; McGill University","funders":"","keywords":"Computer science; Software engineering; Source code; Task (project management); Software; Separation of concerns; Software development; Software system; Code (set theory); Programming language; Systems engineering","score_opus":0.10964718519581704,"score_gpt":0.3609016914329558,"score_spread":0.25125450623713874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153546999","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008804573,0.00032031053,0.9702351,0.00047754883,0.00008934473,0.00020829869,0.0014819748,0.010506837,0.007875996],"genre_scores_gemma":[0.11992464,0.00077325944,0.86055046,0.00022102708,0.000068039044,0.00036950386,0.0058035823,0.004059004,0.008230489],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.996808,0.0010572473,0.00034718262,0.00047552818,0.001110387,0.00020168551],"domain_scores_gemma":[0.9896856,0.004874286,0.00096520985,0.002587159,0.0017113094,0.00017642848],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034129526,0.001178576,0.0004223899,0.0036961983,0.0011753131,0.005305973,0.002091182,0.0024522287,0.0067003174],"category_scores_gemma":[0.016814869,0.001208891,0.00142574,0.0030115915,0.001545534,0.0055382815,0.0033715987,0.0020657168,0.002726634],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002943833,0.0002505311,0.007780003,0.0013211223,0.00013880509,0.002141246,0.011407663,0.073358044,0.019070683,0.5377777,0.02912435,0.3173354],"study_design_scores_gemma":[0.00009328977,0.00011811751,0.0018669942,0.0005797774,0.00015650773,0.0011956012,0.0009828404,0.21352082,0.024541702,0.3052825,0.45151332,0.00014847574],"about_ca_topic_score_codex":0.005438718,"about_ca_topic_score_gemma":0.0044760844,"teacher_disagreement_score":0.0067003174,"about_ca_system_score_codex":0.001470471,"about_ca_system_score_gemma":0.0022767792,"threshold_uncertainty_score":0.022414744},"labels":[],"label_agreement":null},{"id":"W2153678464","doi":"10.1145/2695664.2695900","title":"Design pattern detection using FINDER","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Java; Scripting language; Software design pattern; Design pattern; Open source; Pattern detection; Artificial intelligence; Programming language; Software","score_opus":0.12621350174613832,"score_gpt":0.30270073522245844,"score_spread":0.17648723347632012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153678464","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032471955,0.0010894854,0.74843013,0.0005180872,0.0001564543,0.00045219087,0.005579907,0.20415896,0.00714287],"genre_scores_gemma":[0.14349,0.0007085343,0.8264933,0.0004523545,0.000072192175,0.00047643864,0.010345654,0.010866234,0.0070951553],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957046,0.00061033043,0.0005202245,0.0011562268,0.00175619,0.00025238394],"domain_scores_gemma":[0.98492974,0.009083886,0.0020984868,0.002359061,0.0013335791,0.00019523824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030430993,0.0023183145,0.0013646784,0.008926207,0.00088177045,0.00276108,0.0024603163,0.0017688897,0.012162743],"category_scores_gemma":[0.01673757,0.0014173554,0.0017776086,0.0027275847,0.0009430035,0.0044612503,0.0022379446,0.0013861676,0.005371009],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000589812,0.00032238793,0.019738527,0.0027243712,0.0003613498,0.0018458733,0.001623286,0.0092873145,0.041226037,0.017346248,0.06363788,0.841297],"study_design_scores_gemma":[0.00045755756,0.0006812551,0.015266702,0.0012510306,0.00049751834,0.008253899,0.001175341,0.28401017,0.22974522,0.05665249,0.40150747,0.00050127367],"about_ca_topic_score_codex":0.0019569825,"about_ca_topic_score_gemma":0.0033099663,"teacher_disagreement_score":0.012162743,"about_ca_system_score_codex":0.0008926747,"about_ca_system_score_gemma":0.0021206622,"threshold_uncertainty_score":0.040688455},"labels":[],"label_agreement":null},{"id":"W2154126868","doi":"10.4304/jsw.3.5.26-39","title":"Change Prediction in Object-Oriented Software Systems: A Probabilistic Approach","year":2008,"lang":"en","type":"article","venue":"Journal of Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software system; Software sizing; Software development; Software quality; Reverse engineering; Software metric; Source code; Software maintenance; Software; Software construction; Software engineering; Reliability engineering; Programming language","score_opus":0.038791792737955626,"score_gpt":0.24806302882314094,"score_spread":0.2092712360851853,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154126868","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037234105,0.00029149433,0.9611768,0.00014764637,0.000011133305,0.000095217714,0.00009897432,0.00057815755,0.00036650882],"genre_scores_gemma":[0.6675329,0.000591061,0.3299798,0.000077631244,0.00008686174,0.00031294915,0.000487024,0.000121478966,0.0008102398],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99605006,0.0012945692,0.00027085512,0.00065000274,0.0015954498,0.00013907993],"domain_scores_gemma":[0.97703576,0.016773049,0.0029904775,0.0012961339,0.0016948341,0.00020980254],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040549464,0.0009658966,0.0011028254,0.004783537,0.00076068984,0.0014452216,0.0017274307,0.0013558653,0.0006665886],"category_scores_gemma":[0.023664383,0.0010522432,0.0014602497,0.0024994297,0.001019836,0.002694619,0.0011869599,0.0011833248,0.0002788385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013231911,0.00017400834,0.03419764,0.00024033312,0.00021692374,0.00024805858,0.0003386471,0.80756116,0.0039614877,0.008674898,0.0006230643,0.1436315],"study_design_scores_gemma":[0.0000069693447,0.000040774346,0.0030067414,0.000011727746,0.00003226184,0.00008998319,0.000020214951,0.9893283,0.0006823773,0.0064858277,0.00027396393,0.000020847634],"about_ca_topic_score_codex":0.0059219617,"about_ca_topic_score_gemma":0.0065401755,"teacher_disagreement_score":0.0059219617,"about_ca_system_score_codex":0.0008900736,"about_ca_system_score_gemma":0.0010512172,"threshold_uncertainty_score":0.021444857},"labels":[],"label_agreement":null},{"id":"W2154183829","doi":"10.1145/1985441.1985457","title":"Security versus performance bugs","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":180,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software bug; Computer science; Security bug; Software security assurance; Software; Software quality; Quality (philosophy); Work (physics); Computer security; Software development; Information security; Engineering; Operating system; Security service","score_opus":0.04498061121691938,"score_gpt":0.2522606284033671,"score_spread":0.20728001718644773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154183829","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94669485,0.004017561,0.02033056,0.002613899,0.00024571156,0.00024413771,0.0007333067,0.0005168529,0.024603147],"genre_scores_gemma":[0.99031526,0.00066704577,0.0061138193,0.0002548031,0.00007938357,0.00009698704,0.00023443003,0.00017021722,0.002068089],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97685087,0.0067593167,0.0029230213,0.002723332,0.0085253315,0.0022182136],"domain_scores_gemma":[0.74591833,0.14875522,0.0659662,0.012489182,0.020948581,0.005922426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010791463,0.0010018465,0.0007102825,0.004790384,0.0010185203,0.0029534006,0.0011062894,0.0014761756,0.0062843603],"category_scores_gemma":[0.1252915,0.00055507955,0.0010590925,0.0029724285,0.002559814,0.0057198885,0.0029135896,0.0017311103,0.00089523976],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015086764,0.00055283384,0.6765361,0.0024034604,0.0005366879,0.0012275223,0.009289666,0.006958202,0.014839069,0.029258078,0.0064425194,0.2504472],"study_design_scores_gemma":[0.00021916328,0.0027611065,0.89394325,0.0011818305,0.00050105294,0.0058352426,0.008848143,0.017110135,0.010337446,0.031914156,0.027127098,0.0002214128],"about_ca_topic_score_codex":0.0027273875,"about_ca_topic_score_gemma":0.0027605323,"teacher_disagreement_score":0.010791463,"about_ca_system_score_codex":0.0016909795,"about_ca_system_score_gemma":0.0012463981,"threshold_uncertainty_score":0.057071388},"labels":[],"label_agreement":null},{"id":"W2154234176","doi":"10.1109/spcon.1994.344417","title":"Elicit: a method for eliciting process models","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Process (computing); Dependency (UML); Software engineering; Software; Scale (ratio); Reverse engineering; Product (mathematics); Programming language; Mathematics","score_opus":0.07048182757300119,"score_gpt":0.34525063868672695,"score_spread":0.27476881111372575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154234176","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002826967,0.000026283533,0.99575377,0.00007952055,0.000018384804,0.00018742739,0.00032238552,0.0025277967,0.0008018219],"genre_scores_gemma":[0.009621866,0.00012549735,0.9832681,0.00011684838,0.00002471467,0.0012715618,0.0014665949,0.00092827703,0.0031766247],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9860822,0.00620732,0.0013732119,0.0016189149,0.0044027115,0.00031562327],"domain_scores_gemma":[0.977986,0.014382393,0.0011609915,0.0039507383,0.0021773337,0.00034257743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0085516,0.0025842905,0.0012002416,0.0029274696,0.0015430216,0.0034972946,0.0027943088,0.0032691907,0.018377025],"category_scores_gemma":[0.028271258,0.001882065,0.002378308,0.0025860292,0.0014810764,0.006501457,0.0047782348,0.0042432854,0.008883908],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007036265,0.0003654614,0.0019696557,0.002795249,0.00023459338,0.0008608068,0.005072016,0.023565529,0.029454064,0.23774533,0.055846836,0.6413868],"study_design_scores_gemma":[0.0003620605,0.00027809438,0.0006829062,0.0006611898,0.00016102743,0.001772648,0.00093764893,0.22548845,0.035468154,0.20078547,0.53312415,0.00027830535],"about_ca_topic_score_codex":0.0018318127,"about_ca_topic_score_gemma":0.003263141,"teacher_disagreement_score":0.018377025,"about_ca_system_score_codex":0.0013916271,"about_ca_system_score_gemma":0.0031802305,"threshold_uncertainty_score":0.061477304},"labels":[],"label_agreement":null},{"id":"W2154362664","doi":"10.1007/s10270-015-0483-z","title":"The shape of feature code: an analysis of twenty C-preprocessor-based systems","year":2015,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Conselho Nacional de Desenvolvimento Científico e Tecnológico","keywords":"Feature (linguistics); Computer science; Preprocessor; Outlier; Benchmark (surveying); Program comprehension; Code (set theory); Feature extraction; Data mining; Nesting (process); Pattern recognition (psychology); Artificial intelligence; Programming language; Set (abstract data type); Software","score_opus":0.05149032838202286,"score_gpt":0.29816678618564557,"score_spread":0.2466764578036227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154362664","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99003184,0.00010622269,0.006049014,0.000041703337,0.0000050053427,0.000028598748,0.00029427843,0.0003063305,0.0031371007],"genre_scores_gemma":[0.99465346,0.000040023275,0.0038853833,0.000012115473,0.0000023371865,0.000013205829,0.0005448728,0.00011604944,0.0007325508],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9994369,0.000055902405,0.000024289879,0.00008614139,0.000322062,0.000074632095],"domain_scores_gemma":[0.99374545,0.0029939548,0.00054372754,0.00061823917,0.0018506538,0.00024792404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00041509292,0.00022615216,0.00020103045,0.0014996884,0.00044265648,0.0006484267,0.00051533495,0.0003178077,0.0017257993],"category_scores_gemma":[0.0061160326,0.00018306191,0.0003891777,0.0019092188,0.0005215384,0.00051533297,0.00037369016,0.00032220205,0.00030080442],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021389627,0.0004540708,0.33310705,0.00034747267,0.00011684553,0.0012229556,0.0025338975,0.09900551,0.07499561,0.008446667,0.004805861,0.47282502],"study_design_scores_gemma":[0.000056039164,0.001041157,0.476669,0.000040438048,0.000136027,0.0009144963,0.0014388987,0.46616626,0.04149496,0.0048691086,0.007081076,0.0000924313],"about_ca_topic_score_codex":0.0052482854,"about_ca_topic_score_gemma":0.0058115865,"teacher_disagreement_score":0.0052482854,"about_ca_system_score_codex":0.00085688353,"about_ca_system_score_gemma":0.0008230661,"threshold_uncertainty_score":0.010435462},"labels":[],"label_agreement":null},{"id":"W2154698230","doi":"10.1109/metrics.2005.54","title":"Visualizing Historical Data Using Spectrographs","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Software; Data science; Visualization; Data visualization; Software development; Software engineering; Data mining; Software construction; Programming language","score_opus":0.11074019888734271,"score_gpt":0.35695320302076056,"score_spread":0.24621300413341785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154698230","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21519163,0.003971656,0.69697464,0.0019900384,0.00048685927,0.00034576343,0.03350072,0.02804409,0.019494686],"genre_scores_gemma":[0.5423208,0.0027365717,0.4349618,0.00018815734,0.00026156838,0.00041824786,0.014983683,0.0017049563,0.0024242403],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992442,0.00023219606,0.00009509336,0.00014510864,0.00022832322,0.00005509263],"domain_scores_gemma":[0.99282175,0.003759485,0.001052648,0.0008380852,0.0012426702,0.00028536667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019465587,0.00081357895,0.0005101252,0.015484432,0.00071217655,0.0027944152,0.0005498522,0.00082970463,0.004137609],"category_scores_gemma":[0.0073673245,0.00039560688,0.0005348649,0.011033148,0.00042854247,0.0027813246,0.0015186346,0.0010547352,0.00070675573],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093298557,0.00031128703,0.07292241,0.0021054691,0.0006653934,0.0015317233,0.017333524,0.043194424,0.043911677,0.037101474,0.04989373,0.7300959],"study_design_scores_gemma":[0.0001703945,0.00035706765,0.2017548,0.0011044991,0.0005410872,0.0028550734,0.012680094,0.30137378,0.041733455,0.08463978,0.35216913,0.00062092446],"about_ca_topic_score_codex":0.0046460098,"about_ca_topic_score_gemma":0.0046569775,"teacher_disagreement_score":0.015484432,"about_ca_system_score_codex":0.00054574275,"about_ca_system_score_gemma":0.0005985122,"threshold_uncertainty_score":0.013841689},"labels":[],"label_agreement":null},{"id":"W2154938539","doi":"10.1109/wcre.2001.957820","title":"Maximizing functional cohesion of comprehension environments by integrating user and task knowledge","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Program comprehension; Computer science; Cohesion (chemistry); Comprehension; Human–computer interaction; Task (project management); Information overload; Task analysis; Variety (cybernetics); Software; Abstraction; Software engineering; Artificial intelligence; Software system; World Wide Web; Programming language; Systems engineering; Engineering","score_opus":0.024710804187499356,"score_gpt":0.22468799200094827,"score_spread":0.19997718781344892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154938539","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07776745,0.00016325118,0.9110281,0.0002531333,0.000018632913,0.00019733378,0.000027905715,0.0073677725,0.0031765022],"genre_scores_gemma":[0.3767981,0.00017600106,0.61567926,0.00018969172,0.000057296173,0.0005103108,0.00030719102,0.0031782053,0.0031039754],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9915658,0.003975956,0.0006668389,0.0011830993,0.0020901354,0.0005182922],"domain_scores_gemma":[0.9587223,0.026334763,0.0032762415,0.006520507,0.003913418,0.0012328023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008527571,0.0019210586,0.0012684561,0.0018487638,0.0008996864,0.0040261117,0.0025656607,0.0017218785,0.0024012362],"category_scores_gemma":[0.04371317,0.0014531382,0.0009331635,0.0007049768,0.001261512,0.006465716,0.005505197,0.0017998719,0.0011079342],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015154764,0.002438193,0.012188594,0.001103051,0.00019748947,0.00092746184,0.017501744,0.035214443,0.20349988,0.026472455,0.005640209,0.693301],"study_design_scores_gemma":[0.0007093228,0.0039267475,0.024856512,0.000574334,0.0010678929,0.002072709,0.0055178856,0.5131614,0.25343433,0.11979479,0.074211314,0.00067271054],"about_ca_topic_score_codex":0.00032234707,"about_ca_topic_score_gemma":0.0005709067,"teacher_disagreement_score":0.008527571,"about_ca_system_score_codex":0.00054606015,"about_ca_system_score_gemma":0.001204237,"threshold_uncertainty_score":0.045098662},"labels":[],"label_agreement":null},{"id":"W2154961375","doi":"10.1109/icst.2009.8","title":"Test Redundancy Measurement Based on Coverage Information: Evaluations and Lessons Learned","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Redundancy (engineering); Computer science; Test suite; Reliability engineering; Code coverage; Java; Test case; Data mining; Software; Machine learning; Engineering; Programming language; Operating system","score_opus":0.07346214667201528,"score_gpt":0.33338081168691097,"score_spread":0.2599186650148957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154961375","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.68350625,0.005593547,0.30036157,0.003287756,0.00008408375,0.0002287716,0.00029170758,0.0014647171,0.0051815812],"genre_scores_gemma":[0.92176,0.000864051,0.076466225,0.000115955125,0.00007323024,0.00007367548,0.0002455416,0.00016878147,0.00023265956],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9525915,0.025757326,0.0019344266,0.0025999788,0.016411323,0.0007054604],"domain_scores_gemma":[0.7379544,0.21113649,0.00861351,0.016117703,0.025024822,0.0011530211],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02365936,0.0015141923,0.0014153233,0.0034046264,0.0002916399,0.0016695847,0.0028334726,0.0015250411,0.00082581444],"category_scores_gemma":[0.14741288,0.00047273454,0.00061524403,0.002425041,0.0017368458,0.0041160407,0.001101724,0.0013936092,0.00022425978],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013492801,0.0013864891,0.096729055,0.0014294598,0.00043662402,0.00029045183,0.001163019,0.11783843,0.031061372,0.009639482,0.0033696108,0.7353067],"study_design_scores_gemma":[0.000235929,0.004816057,0.07174348,0.00044353868,0.0003304377,0.0007165308,0.00081923854,0.81863534,0.0877557,0.011161201,0.003116222,0.00022635685],"about_ca_topic_score_codex":0.0032822562,"about_ca_topic_score_gemma":0.0024609917,"teacher_disagreement_score":0.02365936,"about_ca_system_score_codex":0.0019461206,"about_ca_system_score_gemma":0.00085281936,"threshold_uncertainty_score":0.12512416},"labels":[],"label_agreement":null},{"id":"W2155151396","doi":"10.1109/ictai.2004.69","title":"Human perception of software complexity: knowledge discovery from software data","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Institute for Biodiagnostics; University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software; Programming complexity; Perception; Process (computing); Data science; Software development; Software engineering; Data mining; Software construction; Programming language","score_opus":0.08841172534694995,"score_gpt":0.3356505208114203,"score_spread":0.24723879546447036,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155151396","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46749243,0.002791001,0.5141737,0.0054705944,0.00006589652,0.0005491003,0.0019387755,0.0005834123,0.0069351178],"genre_scores_gemma":[0.78982246,0.0015811367,0.20569801,0.00033256313,0.000056591693,0.00031859847,0.0016127971,0.00003910284,0.0005387032],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99254394,0.003371877,0.0005134264,0.0009532505,0.0024300932,0.00018755406],"domain_scores_gemma":[0.91956013,0.06781146,0.0024931359,0.0064483527,0.0031620248,0.00052503165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007528249,0.00050788606,0.00097508234,0.0070896614,0.0008673072,0.004377718,0.0018629353,0.0017331046,0.00093037984],"category_scores_gemma":[0.07046334,0.00059356017,0.0008165074,0.005072783,0.0031853137,0.009170433,0.0024870797,0.0019203869,0.00034630462],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005651193,0.0011371257,0.0817298,0.0021960684,0.00047741982,0.0013362146,0.025490023,0.01624088,0.008786025,0.04085813,0.004360942,0.81682223],"study_design_scores_gemma":[0.00017240364,0.0006816908,0.08690353,0.0013451949,0.0004250152,0.0024260406,0.02564237,0.28626654,0.023319317,0.543642,0.028699802,0.00047609708],"about_ca_topic_score_codex":0.0026308605,"about_ca_topic_score_gemma":0.0026247692,"teacher_disagreement_score":0.007528249,"about_ca_system_score_codex":0.00097365095,"about_ca_system_score_gemma":0.001686046,"threshold_uncertainty_score":0.039813638},"labels":[],"label_agreement":null},{"id":"W2155295509","doi":"10.1145/1368088.1368220","title":"Dynamic round-trip GUI maintenance","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Graphical user interface; Plug-in; Source code; Programming language; Software maintenance; Graphical user interface testing; Context (archaeology); Software engineering; Object-oriented programming; User interface; Code (set theory); Software; Human–computer interaction; Open source; Software development; User interface design; Set (abstract data type)","score_opus":0.017194171549366005,"score_gpt":0.252428089533669,"score_spread":0.23523391798430301,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155295509","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08017219,0.0005092338,0.86592084,0.00047661507,0.00014248544,0.0003117842,0.00028827187,0.04366603,0.008512472],"genre_scores_gemma":[0.5576941,0.00029653427,0.42279142,0.00041543087,0.00008583331,0.00036897213,0.0012991257,0.006170023,0.010878647],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946636,0.0011213728,0.0003440646,0.0013844043,0.0019534952,0.00053305045],"domain_scores_gemma":[0.9667194,0.008272996,0.002276337,0.018612226,0.0034887295,0.0006303332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051764925,0.0011278362,0.0008300142,0.002084615,0.0008324514,0.002516075,0.0047011087,0.0012763919,0.005339699],"category_scores_gemma":[0.023341337,0.0011700239,0.00082841975,0.0010081341,0.001253316,0.0046221344,0.004481102,0.0021254283,0.0021889908],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089364994,0.0005052214,0.014820507,0.00061669457,0.00021437005,0.0013949731,0.0033699085,0.025370156,0.08826347,0.026077883,0.015793988,0.8226792],"study_design_scores_gemma":[0.00034959032,0.0010788236,0.021127105,0.0003830045,0.0005330077,0.005240366,0.0012464143,0.529009,0.24297792,0.06117015,0.1364765,0.00040814848],"about_ca_topic_score_codex":0.0015895679,"about_ca_topic_score_gemma":0.0018940385,"teacher_disagreement_score":0.005339699,"about_ca_system_score_codex":0.00075908296,"about_ca_system_score_gemma":0.0012713022,"threshold_uncertainty_score":0.027376294},"labels":[],"label_agreement":null},{"id":"W2155310431","doi":"10.1109/ccece.2003.1226023","title":"Automating transition from use-cases to class model","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Class (philosophy); Computer science; Transition (genetics); Programming language; Theoretical computer science; Artificial intelligence","score_opus":0.043963058093232664,"score_gpt":0.27460590684335656,"score_spread":0.2306428487501239,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155310431","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019292064,0.00007669916,0.9597945,0.00032507486,0.00004036927,0.0006755161,0.00025833515,0.016884264,0.0026530938],"genre_scores_gemma":[0.18571648,0.00021392525,0.8043028,0.00027506688,0.000042596756,0.00096736324,0.0015562096,0.0034241246,0.0035014045],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98438513,0.0056001176,0.0013338983,0.0020931002,0.0056479084,0.0009397865],"domain_scores_gemma":[0.9616999,0.022644281,0.0021978263,0.010041399,0.0029241566,0.000492362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010798753,0.0013658433,0.001409959,0.0034978662,0.001089591,0.004742397,0.0035013368,0.0020683047,0.0044699432],"category_scores_gemma":[0.05635224,0.0017425249,0.002263863,0.0014303079,0.0014587875,0.0049548154,0.0044583795,0.0030228328,0.0029081837],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006177736,0.00095921836,0.011912123,0.0009000549,0.0002098556,0.0020587074,0.0063741487,0.03956185,0.06571223,0.051115226,0.0149272075,0.8056516],"study_design_scores_gemma":[0.0002026471,0.00028267043,0.005116919,0.00042781662,0.0003023918,0.0018580724,0.0011474763,0.62069273,0.20474246,0.06808973,0.09679794,0.00033907872],"about_ca_topic_score_codex":0.0034390108,"about_ca_topic_score_gemma":0.003301798,"teacher_disagreement_score":0.010798753,"about_ca_system_score_codex":0.0019006702,"about_ca_system_score_gemma":0.0029112939,"threshold_uncertainty_score":0.057109952},"labels":[],"label_agreement":null},{"id":"W2155452486","doi":"10.1109/suite.2009.5070023","title":"Working with search results","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Task (project management); Software; Code (set theory); Source code; Information retrieval; World Wide Web; Data science; Software engineering; Human–computer interaction; Programming language","score_opus":0.04040692068894427,"score_gpt":0.2748433146646666,"score_spread":0.23443639397572233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155452486","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34916937,0.010787002,0.49147746,0.017669268,0.0006598692,0.0020161767,0.004140453,0.019254785,0.104825586],"genre_scores_gemma":[0.55448747,0.0040842155,0.40888003,0.0021668237,0.00043033937,0.0007018049,0.0070632715,0.004925968,0.01726003],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9399148,0.030952433,0.00583898,0.003378031,0.01814104,0.0017747611],"domain_scores_gemma":[0.75737005,0.19414312,0.008722057,0.019203654,0.01828142,0.0022797694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035282385,0.0018469671,0.0027549788,0.013574072,0.0045881867,0.016652556,0.0032727772,0.0035197006,0.013024155],"category_scores_gemma":[0.2306694,0.0014405141,0.0021447458,0.011416095,0.002913961,0.030420527,0.007765161,0.0030864433,0.007027307],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015263667,0.0006318407,0.038864054,0.006564108,0.00064772787,0.003681334,0.20565487,0.0034809124,0.018766155,0.04436092,0.057720315,0.6181014],"study_design_scores_gemma":[0.00048219357,0.0015358579,0.03155184,0.0072475174,0.0019436654,0.011679545,0.19836023,0.048955,0.038507447,0.15805848,0.5004516,0.0012266496],"about_ca_topic_score_codex":0.0042987326,"about_ca_topic_score_gemma":0.0043789092,"teacher_disagreement_score":0.035282385,"about_ca_system_score_codex":0.0019445857,"about_ca_system_score_gemma":0.004384191,"threshold_uncertainty_score":0.18659335},"labels":[],"label_agreement":null},{"id":"W2155534952","doi":"10.1109/icpc.2009.5090066","title":"A bug you like: A framework for automated assignment of bugs","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Heuristics; Task (project management); Software bug; Workload; Set (abstract data type); Software engineering; Space (punctuation); Software; Programming language; Operating system; Systems engineering; Engineering","score_opus":0.018082397108562433,"score_gpt":0.30536034754098995,"score_spread":0.28727795043242754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155534952","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00095279014,0.000076062184,0.99617445,0.00024313218,0.00001599139,0.00010220762,0.00012944451,0.0018696206,0.0004363535],"genre_scores_gemma":[0.037703436,0.00009959972,0.9602971,0.00010272333,0.000041128867,0.00025667134,0.000392677,0.00019748787,0.000909154],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99486357,0.0020572562,0.00034551413,0.0012351231,0.0012070289,0.0002914717],"domain_scores_gemma":[0.9904867,0.004857079,0.0011545611,0.0015597831,0.0013580935,0.0005838353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008343827,0.0016966076,0.0014934295,0.0044552926,0.0016479697,0.0030708835,0.004294659,0.002396944,0.0054483004],"category_scores_gemma":[0.02409463,0.0012954929,0.002109418,0.0028820275,0.0018937836,0.0038894946,0.0031181907,0.0031350458,0.0023786372],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041773796,0.00071229856,0.0057620713,0.0005168746,0.00031671164,0.0004990651,0.001933168,0.21352153,0.006769873,0.16702437,0.031335026,0.57119125],"study_design_scores_gemma":[0.000074580035,0.00009200408,0.0007499821,0.00006448667,0.000053933367,0.00020008125,0.00011765289,0.87478095,0.002313546,0.109580465,0.011897364,0.000074944364],"about_ca_topic_score_codex":0.018188478,"about_ca_topic_score_gemma":0.023831174,"teacher_disagreement_score":0.018188478,"about_ca_system_score_codex":0.0020540084,"about_ca_system_score_gemma":0.0037304652,"threshold_uncertainty_score":0.044126928},"labels":[],"label_agreement":null},{"id":"W2155581635","doi":"10.1109/compsac.2009.38","title":"Predicting Change Impact in Object-Oriented Applications with Bayesian Networks","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Université de Montréal","funders":"","keywords":"Bayesian network; Computer science; Probabilistic logic; Change impact analysis; Graphical model; Causality (physics); Data mining; Bayesian probability; Machine learning; Object-oriented programming; Variable-order Bayesian network; Dynamic Bayesian network; Statistical model; Software; Object (grammar); Artificial intelligence; Bayesian inference; Programming language","score_opus":0.012321824727140041,"score_gpt":0.2726659813105891,"score_spread":0.26034415658344906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155581635","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34918222,0.00044750515,0.64723545,0.0004943091,0.000023844586,0.00013741622,0.00028344302,0.00035519677,0.001840629],"genre_scores_gemma":[0.94911414,0.00036482487,0.04931774,0.00004065212,0.000027664155,0.00009294596,0.00033614083,0.00002794291,0.0006778007],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981523,0.0009957874,0.000088350615,0.00023661311,0.00042182175,0.0001050918],"domain_scores_gemma":[0.98635536,0.011175195,0.0012317786,0.00045027252,0.0006009657,0.00018641843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041061076,0.0008771653,0.0007491648,0.0032428608,0.00042176724,0.0015456739,0.0011096821,0.0011835609,0.0012059544],"category_scores_gemma":[0.022121973,0.0006644603,0.000879222,0.0019799052,0.00060245773,0.003118168,0.00082479336,0.0010534474,0.00020767665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015025122,0.00010888527,0.019105097,0.000047594545,0.00008980847,0.000067567285,0.00010105356,0.94701564,0.00061645923,0.006125963,0.00016755372,0.026404135],"study_design_scores_gemma":[0.0000050739936,0.000023437711,0.0031925163,0.0000067108745,0.000018943763,0.000013401866,0.000018962704,0.98765534,0.00021512812,0.008720558,0.000120475976,0.000009447844],"about_ca_topic_score_codex":0.007873883,"about_ca_topic_score_gemma":0.007046108,"teacher_disagreement_score":0.007873883,"about_ca_system_score_codex":0.001254981,"about_ca_system_score_gemma":0.0005034236,"threshold_uncertainty_score":0.021715403},"labels":[],"label_agreement":null},{"id":"W2155586449","doi":"10.1109/ccece.2002.1013024","title":"A quality assessment model for Java code","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Regina","funders":"","keywords":"Computer science; Java; Reusability; Quality (philosophy); Code (set theory); Interoperability; Software quality; JavaScript; Relation (database); Data mining; Programming language; Operating system; Software development; Software; Set (abstract data type)","score_opus":0.1175190920541189,"score_gpt":0.4037047127958799,"score_spread":0.286185620741761,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155586449","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007963097,0.00017224485,0.9827011,0.0006987469,0.000024928644,0.00023131682,0.00036136774,0.00072064425,0.0071265143],"genre_scores_gemma":[0.30923104,0.00054066477,0.67711973,0.00022508531,0.000062291554,0.0017257027,0.0015215097,0.00025229008,0.009321754],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9944501,0.0014942932,0.0005640899,0.0007160708,0.0025399982,0.00023546001],"domain_scores_gemma":[0.9868548,0.0063911662,0.001612793,0.0008899132,0.0040553226,0.00019602191],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006343558,0.0010423021,0.00042747782,0.0036247254,0.0006206551,0.0034267823,0.0018625541,0.0015426625,0.004205188],"category_scores_gemma":[0.023071121,0.00039228008,0.001490708,0.001985563,0.0011868728,0.0045941514,0.0011539772,0.001349176,0.001748226],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012992873,0.0002279711,0.006399184,0.00041058526,0.00010495491,0.0002827719,0.0012990611,0.40994698,0.0038334185,0.40195507,0.0067998455,0.16861027],"study_design_scores_gemma":[0.000024032179,0.00012540357,0.0012997994,0.00010323901,0.000037260834,0.00012837315,0.00009700702,0.8793425,0.0007635004,0.10674519,0.011298164,0.000035534198],"about_ca_topic_score_codex":0.0069972686,"about_ca_topic_score_gemma":0.005211548,"teacher_disagreement_score":0.0069972686,"about_ca_system_score_codex":0.0029385793,"about_ca_system_score_gemma":0.0026332976,"threshold_uncertainty_score":0.033548355},"labels":[],"label_agreement":null},{"id":"W2155611596","doi":"10.1109/iccl.1992.185464","title":"Static analysis of PostScript code","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Abstract interpretation; Programming language; Computer science; Notation; Code (set theory); Static analysis; Stack (abstract data type); Interpretation (philosophy); Program analysis; Type inference; Theoretical computer science; Artificial intelligence; Arithmetic; Mathematics; Inference","score_opus":0.020369745338002253,"score_gpt":0.27989352733574924,"score_spread":0.259523781997747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155611596","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08761711,0.0003213824,0.8638614,0.00021489563,0.00019562813,0.0002303788,0.0017319523,0.03132372,0.014503568],"genre_scores_gemma":[0.5519386,0.0006520627,0.39580157,0.00027208694,0.00023112076,0.00041443162,0.00673013,0.008941281,0.03501875],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985618,0.00014247236,0.00010581791,0.00022722204,0.0007748966,0.00018775313],"domain_scores_gemma":[0.9965495,0.001002257,0.00045923836,0.00072967506,0.0012042187,0.00005504192],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000870501,0.0007762401,0.00041984755,0.0016514567,0.000659532,0.0012723658,0.0007232388,0.00034997435,0.010441646],"category_scores_gemma":[0.004584399,0.0003535306,0.00071029185,0.0010331967,0.00082553556,0.0016080085,0.00073964696,0.0006001759,0.0025788478],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013729499,0.00017973881,0.014383765,0.0011951562,0.000114786046,0.0028730675,0.0014858876,0.04082612,0.23316583,0.1281604,0.02511548,0.55112684],"study_design_scores_gemma":[0.00005638899,0.00037922626,0.008269968,0.0002723133,0.00015607438,0.0015366208,0.00029191255,0.32358733,0.52243745,0.06295627,0.079923615,0.00013278658],"about_ca_topic_score_codex":0.001767209,"about_ca_topic_score_gemma":0.0019010024,"teacher_disagreement_score":0.010441646,"about_ca_system_score_codex":0.0008140594,"about_ca_system_score_gemma":0.0012798309,"threshold_uncertainty_score":0.034930706},"labels":[],"label_agreement":null},{"id":"W2155943924","doi":"10.1109/icse.2007.80","title":"Suade: Topology-Based Searches for Software Investigation","year":2007,"lang":"en","type":"article","venue":"Proceedings/Proceedings - International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Eclipse; Computer science; Task (project management); Context (archaeology); Software; Software engineering; Software system; Software construction; Source code; Software sizing; Software development; Software evolution; Component-based software engineering; Programming language; Systems engineering; Engineering","score_opus":0.0549667653974837,"score_gpt":0.3028834859539064,"score_spread":0.24791672055642272,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155943924","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016223105,0.0007881923,0.8558192,0.00026612822,0.00008453646,0.00042599783,0.003091111,0.119352944,0.003948781],"genre_scores_gemma":[0.07601379,0.0003657836,0.9062169,0.00011967378,0.00003907355,0.00046972203,0.007897341,0.0066952724,0.0021824187],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976267,0.0007245456,0.00014872134,0.00045338983,0.00091288384,0.00013369377],"domain_scores_gemma":[0.99038833,0.0066298237,0.0005220524,0.0014018173,0.0007773268,0.00028061948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019869201,0.0026685873,0.002050684,0.0076798885,0.0014392191,0.0024940725,0.003721143,0.0027762146,0.009877915],"category_scores_gemma":[0.019524049,0.0017795931,0.0021103022,0.0038328639,0.0010530956,0.004678271,0.0043767225,0.0015477276,0.006380051],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013079608,0.00056327326,0.010406014,0.0024377597,0.00036951853,0.0010042303,0.002116002,0.04444255,0.027778137,0.027378974,0.09601438,0.7861812],"study_design_scores_gemma":[0.00028513285,0.00025929997,0.002258318,0.00024222417,0.00013499754,0.0009863528,0.0005870891,0.86503386,0.021142647,0.048226245,0.06070892,0.0001349943],"about_ca_topic_score_codex":0.0021088186,"about_ca_topic_score_gemma":0.0070115766,"teacher_disagreement_score":0.009877915,"about_ca_system_score_codex":0.0007171633,"about_ca_system_score_gemma":0.0013244844,"threshold_uncertainty_score":0.033044934},"labels":[],"label_agreement":null},{"id":"W2156015672","doi":"10.1109/vlhcc.2008.4639061","title":"Tool support for working with sets of source code entities","year":2008,"lang":"en","type":"article","venue":"Proceedings/Proceedings -- IEEE Symposium on Visual Languages and Human-Centric Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Set (abstract data type); Code (set theory); Source code; KPI-driven code analysis; Code review; Focus (optics); Programming language; Face (sociological concept); Software engineering; Static program analysis; Software; Software development","score_opus":0.02511721503692811,"score_gpt":0.2960395147722234,"score_spread":0.2709222997352953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156015672","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030772008,0.00021734505,0.88225836,0.00072678964,0.00008396757,0.00048457124,0.0006656834,0.080907606,0.003883634],"genre_scores_gemma":[0.16005176,0.0002598167,0.8295401,0.00039297802,0.00007402706,0.0007104183,0.0021483954,0.0046327827,0.0021897224],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9873541,0.0044527305,0.0018500948,0.001886141,0.0040264716,0.0004304925],"domain_scores_gemma":[0.86613405,0.1027593,0.003765141,0.018263802,0.0071146255,0.0019630832],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016344566,0.0016535203,0.0013209493,0.0052156504,0.0011773794,0.0040494716,0.0046671354,0.002221436,0.0072570415],"category_scores_gemma":[0.06358103,0.0014664982,0.0017823662,0.002682594,0.0014609714,0.009972735,0.0065734503,0.0035163756,0.0020877547],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014598703,0.0010731255,0.011556805,0.0027802247,0.00034377677,0.003413143,0.014004861,0.011049766,0.07931245,0.048176114,0.036910687,0.78991914],"study_design_scores_gemma":[0.0013281564,0.0015553554,0.009507781,0.0026189927,0.0004398946,0.009841949,0.0045166533,0.33094698,0.14264886,0.14095667,0.35464862,0.0009900731],"about_ca_topic_score_codex":0.00078146916,"about_ca_topic_score_gemma":0.0013762251,"teacher_disagreement_score":0.016344566,"about_ca_system_score_codex":0.0006924483,"about_ca_system_score_gemma":0.001965302,"threshold_uncertainty_score":0.08643931},"labels":[],"label_agreement":null},{"id":"W2156237258","doi":"10.1109/csmr.2008.4493302","title":"Trend Analysis and Issue Prediction in Large-Scale Open Source Systems","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Eclipse; Computer science; Popularity; Time series; Open source; Software; Series (stratigraphy); Capability Maturity Model; Scale (ratio); Trend analysis; Quality (philosophy); Real-time computing; Operating system; Machine learning","score_opus":0.016454362818573123,"score_gpt":0.26228030921859025,"score_spread":0.24582594640001712,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156237258","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9445034,0.0004985872,0.051063396,0.00046948076,0.000043621,0.000074147356,0.0011045968,0.00061538204,0.0016273829],"genre_scores_gemma":[0.9880364,0.0002470867,0.00977787,0.000020787793,0.000038970797,0.00005590849,0.0013124783,0.0000433366,0.00046715947],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99816257,0.0004860706,0.00021853829,0.00039302168,0.0006090618,0.00013070343],"domain_scores_gemma":[0.96531636,0.023282265,0.006330439,0.0019095469,0.0027531628,0.00040813605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004718649,0.0005524134,0.0005518576,0.0057387236,0.00037625444,0.0011333319,0.00080544327,0.00087040645,0.0010074299],"category_scores_gemma":[0.03492062,0.00037044784,0.00073237665,0.0053142793,0.00041399396,0.0029347076,0.0006233304,0.0010302244,0.00033699488],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003536159,0.00040750476,0.6465964,0.00027813777,0.00029557096,0.00056010886,0.0013306022,0.15509179,0.0038206615,0.005640716,0.0026009665,0.18302387],"study_design_scores_gemma":[0.000015176744,0.00017743505,0.2500658,0.0000337371,0.000048999915,0.00020387117,0.0005339479,0.73845774,0.0013889605,0.0074483096,0.0015811576,0.000044872148],"about_ca_topic_score_codex":0.006755994,"about_ca_topic_score_gemma":0.004983809,"teacher_disagreement_score":0.006755994,"about_ca_system_score_codex":0.00077712344,"about_ca_system_score_gemma":0.0003437644,"threshold_uncertainty_score":0.024954915},"labels":[],"label_agreement":null},{"id":"W2156448859","doi":"10.1109/ms.2009.161","title":"Recommendation Systems for Software Engineering","year":2009,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":387,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; McGill University","funders":"","keywords":"Software engineering; Social software engineering; Software development; Computer science; Software construction; Personal software process; Software peer review; Reuse; Software system; Package development process; Software Engineering Process Group; Software; World Wide Web; Engineering; Operating system","score_opus":0.0223554942090374,"score_gpt":0.26801504227710155,"score_spread":0.24565954806806414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156448859","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045762635,0.013145571,0.9318718,0.005386975,0.0010714188,0.001142151,0.0031709205,0.01763458,0.022000363],"genre_scores_gemma":[0.09236615,0.013358168,0.84980106,0.0022370147,0.001344524,0.0015424852,0.0091798315,0.0007146875,0.029456073],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9918847,0.0031707222,0.0008056827,0.0014468706,0.0024445602,0.00024745366],"domain_scores_gemma":[0.9852014,0.007010658,0.0006633579,0.0030173345,0.0037883373,0.00031887964],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007635613,0.0016628158,0.0022606805,0.0058111707,0.0018496388,0.004321072,0.0029468378,0.004426533,0.024308901],"category_scores_gemma":[0.033909317,0.0009858782,0.002126625,0.008296144,0.000814896,0.006983774,0.002399367,0.0028252231,0.018096408],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003565415,0.00028542653,0.002332051,0.0011341748,0.00062371016,0.00018345808,0.00022863323,0.025739292,0.0013199059,0.060339704,0.12002227,0.7874349],"study_design_scores_gemma":[0.0004041318,0.00035601167,0.0023864412,0.0006105704,0.0005306425,0.00054648955,0.0002629094,0.4890866,0.002844921,0.19893503,0.30377808,0.00025819012],"about_ca_topic_score_codex":0.017261988,"about_ca_topic_score_gemma":0.019069817,"teacher_disagreement_score":0.024308901,"about_ca_system_score_codex":0.0022955607,"about_ca_system_score_gemma":0.0017250469,"threshold_uncertainty_score":0.08132136},"labels":[],"label_agreement":null},{"id":"W2156477392","doi":"10.1109/wcre.1997.624580","title":"Cliche recognition in legacy software: a scalable, knowledge-based approach","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Cliché; Computer science; Robustness (evolution); Reverse engineering; Scalability; Software; Software system; Software engineering; Software development; Artificial intelligence; Programming language; Database","score_opus":0.05300804821581124,"score_gpt":0.2626545273488347,"score_spread":0.20964647913302348,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156477392","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03158038,0.0004644361,0.9412504,0.0010733246,0.000059959162,0.0005030132,0.00041477682,0.020096537,0.0045571527],"genre_scores_gemma":[0.17272876,0.00037078187,0.82051116,0.00029122256,0.00006494157,0.00023697279,0.00106172,0.0006887678,0.0040455796],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99604756,0.00045936398,0.00023443808,0.00096297974,0.0019916808,0.00030403424],"domain_scores_gemma":[0.9942643,0.0016753152,0.0006334412,0.001745841,0.0013725,0.00030859097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018127925,0.0015754558,0.0026548076,0.0050207498,0.0019643959,0.0051342766,0.006067179,0.0031038474,0.004939492],"category_scores_gemma":[0.008553429,0.0010961838,0.0019565271,0.003526856,0.0018048726,0.008224546,0.005391674,0.002688915,0.0026910154],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002989409,0.000636897,0.0039018465,0.00036244368,0.00012513326,0.00065258064,0.0007061259,0.03668934,0.024065163,0.009915063,0.011713575,0.9109328],"study_design_scores_gemma":[0.00006679405,0.00009471962,0.001683783,0.000036156816,0.000077321754,0.00049291,0.00073301396,0.9442777,0.018015064,0.027803235,0.0066336547,0.000085613014],"about_ca_topic_score_codex":0.011710761,"about_ca_topic_score_gemma":0.01965194,"teacher_disagreement_score":0.011710761,"about_ca_system_score_codex":0.001816747,"about_ca_system_score_gemma":0.0023402192,"threshold_uncertainty_score":0.02328521},"labels":[],"label_agreement":null},{"id":"W2156491770","doi":"10.1109/wcre.2007.52","title":"Visualizing Software Architecture Evolution Using Change-Sets","year":2007,"lang":"en","type":"article","venue":"Proceedings - Working Conference on Reverse Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software evolution; Computer science; Software architecture; Java; Architecture; Software system; Software; Set (abstract data type); Resource-oriented architecture; Software engineering; Multilayered architecture; Software architecture description; Reference architecture; Programming language; Component-based software engineering; Software construction","score_opus":0.06748153835963974,"score_gpt":0.29869475212898533,"score_spread":0.2312132137693456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156491770","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24363261,0.0015412119,0.7008513,0.0018169885,0.00023758855,0.00036663187,0.004221874,0.025071282,0.022260519],"genre_scores_gemma":[0.5940761,0.0007839691,0.3982713,0.00013685245,0.000054204997,0.00028032725,0.0025012617,0.0014014024,0.0024946535],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992349,0.0003115691,0.000059018697,0.0001116082,0.00022213688,0.00006085871],"domain_scores_gemma":[0.9946807,0.0033520819,0.0004928125,0.0006341544,0.0005711161,0.00026913936],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018034384,0.0009245455,0.00044692846,0.006425937,0.00091946125,0.002466597,0.00096058124,0.0013117827,0.0056104246],"category_scores_gemma":[0.0065861833,0.00049570296,0.0008007483,0.0032838616,0.00075379375,0.0032929545,0.0022182697,0.0019805338,0.00059060485],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014691043,0.0007527267,0.0506884,0.001955429,0.000529955,0.0023571374,0.04185494,0.15123056,0.06887048,0.119590074,0.031914566,0.5287866],"study_design_scores_gemma":[0.0002834734,0.00048420957,0.05088471,0.0005787111,0.00030073084,0.0016993505,0.0070629106,0.6563153,0.045339502,0.11696177,0.1195842,0.0005051513],"about_ca_topic_score_codex":0.0057418426,"about_ca_topic_score_gemma":0.00655772,"teacher_disagreement_score":0.006425937,"about_ca_system_score_codex":0.0009209153,"about_ca_system_score_gemma":0.0005609235,"threshold_uncertainty_score":0.018768787},"labels":[],"label_agreement":null},{"id":"W2156555135","doi":"10.1145/1137983.1138008","title":"A lightweight approach to technical risk estimation via probabilistic impact analysis","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Probabilistic logic; Eclipse; Computer science; Risk management; Context (archaeology); Event (particle physics); Risk analysis (engineering); Estimation; New product development; Product (mathematics); Systems engineering; Engineering; Artificial intelligence; Business","score_opus":0.008607746860523506,"score_gpt":0.26405845260903815,"score_spread":0.25545070574851464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156555135","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011647813,0.000029003364,0.9978218,0.000041913754,0.000004758847,0.000019668088,0.0000188538,0.0002377388,0.0006614908],"genre_scores_gemma":[0.15873396,0.0002587482,0.8384838,0.000063138454,0.00007167049,0.00024219099,0.00015247635,0.00021617507,0.0017778744],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99385935,0.0016503267,0.00029995435,0.00054561294,0.0034035943,0.00024110197],"domain_scores_gemma":[0.98332787,0.010065334,0.0016486206,0.002997342,0.0017826833,0.00017812761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00530986,0.0017375586,0.0013303681,0.005280397,0.0010569486,0.0033479708,0.0026455428,0.0015794135,0.0039408803],"category_scores_gemma":[0.029200964,0.0011865082,0.0022181103,0.0025510616,0.0014556174,0.004925024,0.004050974,0.002905667,0.0011738637],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012167303,0.0001729767,0.004843185,0.00031049456,0.00030281657,0.00039833118,0.0005330858,0.4161083,0.011764552,0.20916605,0.0023403186,0.3539382],"study_design_scores_gemma":[0.00001371549,0.00005749292,0.00091641507,0.00006369448,0.0000847464,0.00023033736,0.000048344322,0.833698,0.003397812,0.15836962,0.0030593872,0.000060405335],"about_ca_topic_score_codex":0.0021374382,"about_ca_topic_score_gemma":0.0021584416,"teacher_disagreement_score":0.00530986,"about_ca_system_score_codex":0.0011265031,"about_ca_system_score_gemma":0.0017702105,"threshold_uncertainty_score":0.028081596},"labels":[],"label_agreement":null},{"id":"W2156581466","doi":"10.1109/wse.2003.1234008","title":"Resolution of static clones in dynamic Web pages","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Reuse; Parsing; Cloning (programming); Syntax; World Wide Web; Web application; Programming language; Code (set theory); Source code; clone (Java method); Web page; Code reuse; Software; Resolution (logic); Information retrieval; Artificial intelligence; Engineering; Biology; Genetics","score_opus":0.013414160561565467,"score_gpt":0.26675274508512253,"score_spread":0.25333858452355706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156581466","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38441825,0.0027818347,0.5881969,0.0013675548,0.00034857466,0.0004482177,0.0005966803,0.011263973,0.010577974],"genre_scores_gemma":[0.4595587,0.0009134245,0.5193349,0.0006225644,0.00016519889,0.00018874691,0.0014093329,0.0026861895,0.015120974],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99591947,0.0007223951,0.00039586608,0.0007355616,0.0019382951,0.00028844277],"domain_scores_gemma":[0.96773326,0.010483718,0.005328957,0.008900482,0.0067737103,0.0007797969],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028120955,0.00064942456,0.0008740716,0.004340733,0.0021594781,0.00290392,0.0012326504,0.0024844264,0.0022116504],"category_scores_gemma":[0.025685525,0.0007286123,0.00090068736,0.003997807,0.0012618738,0.0035110076,0.0025864646,0.0017998976,0.0017692811],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041338682,0.00040128504,0.05505759,0.0007472651,0.00016703118,0.008371443,0.010231081,0.005636351,0.16166681,0.028775664,0.011050389,0.7174817],"study_design_scores_gemma":[0.00010993826,0.0005013844,0.06506614,0.0007793266,0.000650102,0.028442742,0.0072544,0.07366545,0.52703357,0.08256212,0.21349333,0.00044156806],"about_ca_topic_score_codex":0.0015940663,"about_ca_topic_score_gemma":0.0022880777,"teacher_disagreement_score":0.004340733,"about_ca_system_score_codex":0.0007474152,"about_ca_system_score_gemma":0.0013909162,"threshold_uncertainty_score":0.0148720145},"labels":[],"label_agreement":null},{"id":"W2156587424","doi":"10.1145/1808901.1808913","title":"Clone detection by exploiting assembler","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Source code; clone (Java method); Code (set theory); Programming language; Parallel computing; Assembly language; Software; Set (abstract data type); Biology","score_opus":0.011207888429076954,"score_gpt":0.249570538985326,"score_spread":0.23836265055624906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156587424","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18289296,0.0012057994,0.7897822,0.00014223863,0.00008290108,0.0001812243,0.0002947393,0.023471376,0.0019464801],"genre_scores_gemma":[0.31672364,0.00062903576,0.67621475,0.00016614211,0.000071379196,0.00012239668,0.0015333196,0.002048932,0.0024904248],"study_design_codex":"bench_or_experimental","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99640334,0.000801072,0.00031636527,0.00066515873,0.0015695436,0.00024457028],"domain_scores_gemma":[0.9903756,0.0036256842,0.0020404481,0.0021827274,0.0015830651,0.00019244719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017705965,0.0012715777,0.0011050538,0.0042616064,0.00068027456,0.0014738427,0.0014225404,0.0011059268,0.001424161],"category_scores_gemma":[0.008590552,0.00057793583,0.0011758264,0.0017041496,0.0006716957,0.002039017,0.0011907445,0.0009342173,0.0012086205],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039853147,0.000221249,0.023338048,0.00051024894,0.00027417,0.0017890058,0.0008975308,0.0073590646,0.6091476,0.002897975,0.0014331003,0.3517335],"study_design_scores_gemma":[0.000042228385,0.00091110857,0.021229666,0.00008913733,0.00027155873,0.0056883227,0.0002645168,0.15937424,0.7926467,0.0031259814,0.0161487,0.0002078217],"about_ca_topic_score_codex":0.0009308126,"about_ca_topic_score_gemma":0.0011473184,"teacher_disagreement_score":0.0042616064,"about_ca_system_score_codex":0.0003442849,"about_ca_system_score_gemma":0.0004110744,"threshold_uncertainty_score":0.009363949},"labels":[],"label_agreement":null},{"id":"W2156672158","doi":"10.1109/acom.2007.4","title":"Identifying, Assigning, and Quantifying Crosscutting Concerns","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Identification (biology); Modularity (biology); Software engineering; Business process reengineering; Ambiguity; Software quality; Code (set theory); Suite; Quality (philosophy); Software; Software development; Programming language; Engineering; Set (abstract data type)","score_opus":0.08655784220837867,"score_gpt":0.37326549733640396,"score_spread":0.2867076551280253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156672158","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56572133,0.00056316896,0.42806432,0.0003329042,0.000034550725,0.00065504794,0.00031902755,0.0019385379,0.0023711654],"genre_scores_gemma":[0.53215444,0.00021277994,0.4655654,0.000076099175,0.000012989987,0.00033052333,0.0004988476,0.00021972887,0.00092918787],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9871579,0.004417957,0.0017112057,0.001625136,0.004727211,0.00036057335],"domain_scores_gemma":[0.9012438,0.043363526,0.022481741,0.008904228,0.022938337,0.0010682663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008418346,0.0013958458,0.0008133713,0.0068703145,0.001152414,0.0018711302,0.0011025497,0.0010632317,0.00063463394],"category_scores_gemma":[0.06985132,0.0006614276,0.00053073966,0.002962185,0.0008092788,0.0030167673,0.0018081435,0.0010323427,0.00021160675],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018256706,0.00047412992,0.33105904,0.0010068273,0.00018565069,0.0005385358,0.0038879188,0.02799825,0.07421508,0.007694285,0.0018546003,0.5509031],"study_design_scores_gemma":[0.00009514489,0.0012543845,0.40814835,0.00048425954,0.000483815,0.0029899406,0.0032016726,0.31899866,0.21791367,0.028117169,0.01798462,0.00032833187],"about_ca_topic_score_codex":0.0032672489,"about_ca_topic_score_gemma":0.007089182,"teacher_disagreement_score":0.008418346,"about_ca_system_score_codex":0.0011688871,"about_ca_system_score_gemma":0.0022640603,"threshold_uncertainty_score":0.044521034},"labels":[],"label_agreement":null},{"id":"W2156833313","doi":"10.1145/1368088.1368151","title":"An approach to detecting duplicate bug reports using natural language and execution information","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":544,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Eclipse; Software bug; Natural language; Security bug; Open source; Information retrieval; Natural language processing; Programming language; Software; Information security; Computer security","score_opus":0.018405413256578004,"score_gpt":0.26399311209562837,"score_spread":0.24558769883905035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156833313","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0997027,0.0014408377,0.87470245,0.0013937949,0.00019407507,0.0014345486,0.0016134356,0.0165717,0.0029465188],"genre_scores_gemma":[0.1645165,0.00028123075,0.8308689,0.00047828016,0.000105894804,0.00044403505,0.0013827601,0.00027903877,0.0016433445],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9853231,0.0034842615,0.0018681416,0.0031210296,0.005830804,0.00037258165],"domain_scores_gemma":[0.9296833,0.031111605,0.015463564,0.008280903,0.014701485,0.00075908453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0090852,0.0019091788,0.0021365648,0.014981568,0.0016913314,0.0025297333,0.0038036704,0.0027102346,0.0011947232],"category_scores_gemma":[0.045965422,0.0012783507,0.0022386878,0.0059842174,0.0015848926,0.0048378417,0.0021877498,0.0020586692,0.0006546177],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087127736,0.0013895893,0.058072593,0.0019953297,0.000660558,0.0034582408,0.0046964437,0.012280146,0.06748276,0.007841917,0.010206024,0.83104515],"study_design_scores_gemma":[0.00075957144,0.0017481643,0.07598227,0.0007097548,0.0018564895,0.017096318,0.003391933,0.6428931,0.16365442,0.0339969,0.05646628,0.0014448548],"about_ca_topic_score_codex":0.011898058,"about_ca_topic_score_gemma":0.014724553,"teacher_disagreement_score":0.014981568,"about_ca_system_score_codex":0.0017916561,"about_ca_system_score_gemma":0.0053588822,"threshold_uncertainty_score":0.04804766},"labels":[],"label_agreement":null},{"id":"W2156949431","doi":"10.1109/icsm.2009.5306295","title":"Software maintainability benefits from annotation-driven code","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Computer science; Java; Suite; Metadata; Maintainability; Code (set theory); Programming language; Software engineering; Software; Annotation; Software maintenance; KPI-driven code analysis; Operating system; Software development; Database; Static program analysis; Artificial intelligence","score_opus":0.01650307929797773,"score_gpt":0.26309105026159746,"score_spread":0.24658797096361973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156949431","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023144605,0.00015746865,0.9509468,0.0009682428,0.00008786527,0.00014159932,0.00018008935,0.017669203,0.0067040627],"genre_scores_gemma":[0.2915144,0.00037780142,0.6802107,0.00053164264,0.0001889463,0.00041316016,0.0018703829,0.014088775,0.010804145],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.990013,0.002218629,0.0008208066,0.0012843359,0.0053439974,0.00031920752],"domain_scores_gemma":[0.9462038,0.017532913,0.003161826,0.02319823,0.009415737,0.00048746078],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008327414,0.0010481742,0.00049967505,0.0018501979,0.00096200674,0.0028818727,0.0029338442,0.0014465675,0.001872659],"category_scores_gemma":[0.047796145,0.0010016665,0.00086735643,0.0014008878,0.00200447,0.005016792,0.00335305,0.002738701,0.0018756838],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038460843,0.0005430644,0.010974973,0.00063898525,0.000108817534,0.001372492,0.005127694,0.04837156,0.104409575,0.12478415,0.013361943,0.6899222],"study_design_scores_gemma":[0.0002541363,0.00048755095,0.008415615,0.0004816326,0.00032862803,0.00244589,0.0007064292,0.34092295,0.20179896,0.22072458,0.2230913,0.00034232155],"about_ca_topic_score_codex":0.0019368044,"about_ca_topic_score_gemma":0.001958077,"teacher_disagreement_score":0.008327414,"about_ca_system_score_codex":0.0008066288,"about_ca_system_score_gemma":0.0020668702,"threshold_uncertainty_score":0.044040084},"labels":[],"label_agreement":null},{"id":"W2156993248","doi":"10.1109/ccece.1999.807217","title":"Rough software deployability control system: design and implementation","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Cyclomatic complexity; Computer science; Table (database); Software; Decision table; Java; Software system; Software deployment; Control system; Data mining; Software engineering; Rough set; Operating system; Engineering","score_opus":0.017613223885161313,"score_gpt":0.2750425874162459,"score_spread":0.2574293635310846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156993248","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0040945825,0.000047865175,0.988613,0.00006308757,0.000027762486,0.00030622014,0.00007846621,0.00574908,0.0010198387],"genre_scores_gemma":[0.1725615,0.00016445425,0.82133216,0.00009068196,0.000051245617,0.0010202188,0.00045731472,0.00045111874,0.003871258],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99858946,0.0002194169,0.00011801296,0.00023642862,0.0007708398,0.0000659454],"domain_scores_gemma":[0.998464,0.0005455018,0.00013092955,0.00022022534,0.00056873803,0.000070665825],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018860518,0.00052324426,0.00094607816,0.001528748,0.0005012125,0.0021367306,0.0018059545,0.00074441015,0.007836505],"category_scores_gemma":[0.0047611743,0.0006392456,0.00058706995,0.00079014944,0.00060137693,0.0010352377,0.0006773353,0.0008393827,0.0018266982],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067890587,0.00033144644,0.0021754864,0.0006211771,0.00015254732,0.00029822913,0.0005907466,0.25069523,0.03959067,0.045254394,0.010283181,0.649328],"study_design_scores_gemma":[0.00013123281,0.0002501141,0.00079525367,0.000040211027,0.000071258924,0.00013783202,0.000048370628,0.95637774,0.020708751,0.005726176,0.015651785,0.000061370476],"about_ca_topic_score_codex":0.0036473107,"about_ca_topic_score_gemma":0.001520591,"teacher_disagreement_score":0.007836505,"about_ca_system_score_codex":0.0007562638,"about_ca_system_score_gemma":0.0010770988,"threshold_uncertainty_score":0.026215672},"labels":[],"label_agreement":null},{"id":"W2157007220","doi":"10.1109/icpc.2007.7","title":"A Hybrid Program Model for Object-Oriented Reverse Engineering","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Reverse engineering; Granularity; Program comprehension; Scalability; Set (abstract data type); Object-oriented programming; Programming language; Software engineering; Object (grammar); Focus (optics); Domain (mathematical analysis); Unified Modeling Language; Software; Theoretical computer science; Artificial intelligence; Software system; Database","score_opus":0.018047304108090127,"score_gpt":0.2830028713124129,"score_spread":0.2649555672043228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157007220","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003522324,0.000054539338,0.9932473,0.00019529129,0.000009656392,0.00006938646,0.000052710475,0.00048190303,0.0023668974],"genre_scores_gemma":[0.10924497,0.00021915819,0.8846783,0.00015823022,0.000022544396,0.00059561554,0.00027301535,0.00026434893,0.0045437426],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99860495,0.0004969755,0.00009418261,0.0002062036,0.0005217849,0.0000759072],"domain_scores_gemma":[0.99807966,0.0007357037,0.00014971264,0.0006402196,0.00032079083,0.00007401734],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014715182,0.0006756075,0.00048979407,0.0012864623,0.000601493,0.0022933588,0.0020021382,0.0017316072,0.0025148008],"category_scores_gemma":[0.0029938265,0.0005339272,0.0012870149,0.0012717812,0.0019322728,0.0040775705,0.0019866743,0.001756305,0.0008517129],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006333721,0.00013538728,0.0009307708,0.00017704723,0.00004143066,0.0004601256,0.001280471,0.10559481,0.0080658,0.82890147,0.0019671768,0.052382104],"study_design_scores_gemma":[0.00005423844,0.000120611614,0.00020199973,0.00006184355,0.000056045228,0.0003668436,0.00016129969,0.6076029,0.0045797634,0.34391537,0.042837474,0.00004162287],"about_ca_topic_score_codex":0.0024593393,"about_ca_topic_score_gemma":0.0029258248,"teacher_disagreement_score":0.0025148008,"about_ca_system_score_codex":0.0009109357,"about_ca_system_score_gemma":0.0013691625,"threshold_uncertainty_score":0.008412898},"labels":[],"label_agreement":null},{"id":"W2157018921","doi":"10.1145/1370175.1370222","title":"Composing knowledge fragments","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Programmer; Software engineering; Knowledge-based systems; Process (computing); Domain knowledge; Software; Knowledge engineering; Knowledge management; Human–computer interaction; Programming language","score_opus":0.03588630637404694,"score_gpt":0.28402878799158127,"score_spread":0.24814248161753433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157018921","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022893485,0.00080455124,0.9302067,0.0012173469,0.00021637997,0.0007949473,0.0006843711,0.003595224,0.03958691],"genre_scores_gemma":[0.15494515,0.00093960506,0.81911355,0.00059065234,0.00015738883,0.0007403052,0.0031206647,0.0013213007,0.019071395],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99231964,0.002058167,0.00063679076,0.0013628924,0.002992289,0.0006302687],"domain_scores_gemma":[0.9785478,0.011916451,0.00066661945,0.005864288,0.0022954152,0.0007094452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050944136,0.0017382966,0.0012950277,0.0035447632,0.00240994,0.005708989,0.0038276752,0.0027297088,0.018529873],"category_scores_gemma":[0.031684343,0.0019130892,0.0028160803,0.003587093,0.00394094,0.015810376,0.01132433,0.0038115412,0.0041239685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058789196,0.00036713685,0.0040770443,0.0015288312,0.00028228018,0.0048092846,0.016389655,0.037944946,0.01862203,0.45148084,0.014900265,0.4490098],"study_design_scores_gemma":[0.00008756183,0.00021388862,0.0013159168,0.0006087039,0.0003865506,0.0013928841,0.0038603412,0.09002357,0.018959822,0.7191666,0.16383623,0.00014798545],"about_ca_topic_score_codex":0.0069723777,"about_ca_topic_score_gemma":0.007055937,"teacher_disagreement_score":0.018529873,"about_ca_system_score_codex":0.0018829502,"about_ca_system_score_gemma":0.0027962974,"threshold_uncertainty_score":0.061988592},"labels":[],"label_agreement":null},{"id":"W2157096094","doi":"10.1109/wcre.2000.891476","title":"A maintainability model for industrial software systems using design level metrics","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nortel (Canada); University of Waterloo","funders":"","keywords":"Maintainability; Software sizing; Software metric; Software construction; Computer science; Verification and validation; Reliability engineering; Software measurement; Software reliability testing; Software development; Software maintenance; Software system; Software engineering; Software; Software design; Engineering; Operating system","score_opus":0.4031369177315825,"score_gpt":0.32489819060362235,"score_spread":0.07823872712796015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157096094","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024863804,0.0001786996,0.97146094,0.00028041122,0.000016117741,0.00008173314,0.0002821449,0.00090371474,0.0019324817],"genre_scores_gemma":[0.6847389,0.00063696457,0.30692685,0.00012485006,0.000071686336,0.0009989657,0.0015312092,0.00019346084,0.0047771707],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999033,0.00032606325,0.000057794223,0.00018903824,0.00033117653,0.000062923296],"domain_scores_gemma":[0.9969614,0.0020232103,0.00034599943,0.00019086813,0.0004297585,0.000048739676],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017014731,0.0009748845,0.0004760601,0.001500129,0.00031872495,0.0010203759,0.0013622651,0.0010855994,0.0016082872],"category_scores_gemma":[0.008096816,0.00033671214,0.0007394644,0.0010777526,0.0005375712,0.0018284372,0.0004621571,0.001052383,0.0006535262],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003827751,0.00012698572,0.0018850223,0.00007721791,0.000034651068,0.00007102598,0.00010838868,0.93090206,0.0021184653,0.019481188,0.001236962,0.043919668],"study_design_scores_gemma":[0.0000049997548,0.000048896054,0.0005440924,0.000007682556,0.0000098010205,0.000016872807,0.000004427485,0.9923787,0.00015671542,0.0063184937,0.00050431053,0.0000050965677],"about_ca_topic_score_codex":0.008932896,"about_ca_topic_score_gemma":0.0063259145,"teacher_disagreement_score":0.008932896,"about_ca_system_score_codex":0.001245291,"about_ca_system_score_gemma":0.0010919897,"threshold_uncertainty_score":0.017761767},"labels":[],"label_agreement":null},{"id":"W2157411243","doi":"10.1109/wpc.2000.852482","title":"Tracing object-oriented code into functional requirements","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Programming language; Tracing; Program comprehension; Java; Identifier; Source code; Software engineering; Software maintenance; Documentation; TRACE (psycholinguistics); Software development; Software system; Software","score_opus":0.04631600846662325,"score_gpt":0.2761821812184072,"score_spread":0.22986617275178395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157411243","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07127367,0.00009241494,0.92043084,0.0005610582,0.000025943871,0.00031657418,0.00018064036,0.002447954,0.0046709436],"genre_scores_gemma":[0.3133295,0.00024205027,0.6800453,0.00015112056,0.000019738778,0.00034045835,0.0005500523,0.00090868154,0.0044130874],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952402,0.0020008879,0.00036162997,0.0005193352,0.0016978392,0.00018006901],"domain_scores_gemma":[0.94247603,0.030607652,0.007893929,0.011747034,0.00672677,0.00054860144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006303316,0.0007739574,0.00025678086,0.0038056003,0.0009584155,0.002478493,0.0012884849,0.0013125545,0.002235344],"category_scores_gemma":[0.05895138,0.0007615816,0.0005362897,0.00156636,0.0019365298,0.0045236796,0.001957289,0.0013376896,0.0010468329],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027230926,0.00073038426,0.029823285,0.0010982606,0.000077851575,0.0015115688,0.03041772,0.046381984,0.05271934,0.25392482,0.0037223368,0.57932013],"study_design_scores_gemma":[0.0001296709,0.00096819573,0.025086517,0.000989329,0.0001371073,0.0018814761,0.00654966,0.3153253,0.13401183,0.40188456,0.11276425,0.00027205795],"about_ca_topic_score_codex":0.004472544,"about_ca_topic_score_gemma":0.0038356234,"teacher_disagreement_score":0.006303316,"about_ca_system_score_codex":0.0014214085,"about_ca_system_score_gemma":0.003561042,"threshold_uncertainty_score":0.033335507},"labels":[],"label_agreement":null},{"id":"W2157458605","doi":"10.1145/1985441.1985483","title":"Apples vs. oranges?","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Eclipse; Java; Computer science; Source code; Code (set theory); Open source; Programming language; Software engineering; Software; Physics; Astronomy","score_opus":0.044638373207198864,"score_gpt":0.24799397067131157,"score_spread":0.2033555974641127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157458605","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7945795,0.0027951035,0.009906881,0.008802648,0.0020341396,0.00022903763,0.001539515,0.0009948647,0.17911828],"genre_scores_gemma":[0.9750123,0.00064059463,0.007350056,0.0029516055,0.00020634865,0.00007603837,0.0005732473,0.0006199851,0.01256992],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99658847,0.0011919884,0.00018769366,0.0007973246,0.00093722774,0.00029725657],"domain_scores_gemma":[0.96683407,0.02556855,0.0024172785,0.0017391563,0.00227607,0.0011649276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050287526,0.00056862953,0.00092636456,0.0012664258,0.00083999167,0.0032249035,0.00058364496,0.0019292608,0.016202902],"category_scores_gemma":[0.033894263,0.0002570917,0.000436721,0.0008737714,0.0010468628,0.006070493,0.0014990168,0.0016736359,0.004130407],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.025897995,0.0027505436,0.201423,0.002612949,0.0007614507,0.0019594997,0.009765195,0.0022731482,0.056675434,0.12646349,0.08389598,0.4855214],"study_design_scores_gemma":[0.0029121793,0.008884858,0.35860166,0.0010236198,0.0011177928,0.0043610726,0.02602695,0.026354901,0.065819904,0.13869137,0.36570174,0.0005039906],"about_ca_topic_score_codex":0.0012816613,"about_ca_topic_score_gemma":0.0016150905,"teacher_disagreement_score":0.016202902,"about_ca_system_score_codex":0.00041009823,"about_ca_system_score_gemma":0.0002728907,"threshold_uncertainty_score":0.054204106},"labels":[],"label_agreement":null},{"id":"W2157731798","doi":"10.1109/wcre.2005.19","title":"Extracting and Representing Cross-Language Dependencies in Diverse Software Systems","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Programming language; Java; Software engineering; Software development; Software system; Reverse engineering; Schema (genetic algorithms); Software; Information retrieval","score_opus":0.015702735266102524,"score_gpt":0.2842401394281578,"score_spread":0.2685374041620553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157731798","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1545664,0.00048499857,0.8371994,0.00037045364,0.000036063127,0.00016931415,0.0011198113,0.003368136,0.002685526],"genre_scores_gemma":[0.2935749,0.0005754245,0.6994567,0.00012829223,0.000020453752,0.00014428164,0.0036941823,0.00076335465,0.0016424421],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982728,0.00045071967,0.00018616744,0.00037027345,0.0006194542,0.000100710335],"domain_scores_gemma":[0.9933427,0.0037888174,0.00077985047,0.0011094221,0.0008939136,0.00008522932],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017775606,0.0007316097,0.0006210855,0.004947439,0.0010337407,0.0020270573,0.0010569164,0.0010966517,0.0013002452],"category_scores_gemma":[0.009260879,0.0006770139,0.0008557429,0.0033586144,0.0007445985,0.0038431867,0.0024063385,0.001361256,0.00036878505],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040819455,0.00036228105,0.045831043,0.0012564647,0.00025287914,0.009801099,0.013057754,0.054846484,0.0714139,0.07079374,0.0062280325,0.7257482],"study_design_scores_gemma":[0.0001173984,0.00034072268,0.03569057,0.00071881304,0.0006063435,0.005641154,0.0052163787,0.5149985,0.16145505,0.14990413,0.12502766,0.00028324482],"about_ca_topic_score_codex":0.003968221,"about_ca_topic_score_gemma":0.008056518,"teacher_disagreement_score":0.004947439,"about_ca_system_score_codex":0.0008032397,"about_ca_system_score_gemma":0.0012148115,"threshold_uncertainty_score":0.009400725},"labels":[],"label_agreement":null},{"id":"W2157766445","doi":"10.1109/ccece.1999.808192","title":"Learning how to program","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Debugging; Computer science; Software engineering; Variety (cybernetics); Programming language; Representation (politics); Plan (archaeology); Software; Code (set theory); Artificial intelligence; Software development","score_opus":0.018481571908912746,"score_gpt":0.2812583978692551,"score_spread":0.26277682596034235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157766445","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024849627,0.0010940148,0.938312,0.007826566,0.00027375304,0.00019187086,0.00057210325,0.0034106325,0.023469333],"genre_scores_gemma":[0.1304951,0.0029683188,0.844154,0.001466265,0.00015879289,0.00023324625,0.001992644,0.000964551,0.017567078],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9987469,0.00037286908,0.00008267277,0.00043946508,0.00028280081,0.000075331285],"domain_scores_gemma":[0.9965463,0.0019851932,0.00018866628,0.0007398625,0.00042427634,0.00011569313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017716417,0.0010427878,0.0004130587,0.00070636085,0.0007786619,0.0032745618,0.0017473964,0.0013947125,0.011901847],"category_scores_gemma":[0.01104397,0.0006594493,0.0011176671,0.00046325885,0.00226436,0.009382661,0.0015675677,0.0031641014,0.0046170983],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000092413975,0.00025694404,0.005109597,0.0006914675,0.000087521745,0.00029443105,0.0018608101,0.016464077,0.008682525,0.3031998,0.034083523,0.629177],"study_design_scores_gemma":[0.000047173122,0.0001291062,0.0014375442,0.000426397,0.00008212925,0.0005607588,0.0008483379,0.09948489,0.0131170675,0.703461,0.18033536,0.000070088805],"about_ca_topic_score_codex":0.002433561,"about_ca_topic_score_gemma":0.003981052,"teacher_disagreement_score":0.011901847,"about_ca_system_score_codex":0.00096971134,"about_ca_system_score_gemma":0.0022184364,"threshold_uncertainty_score":0.039815664},"labels":[],"label_agreement":null},{"id":"W2158046712","doi":"10.1109/coginf.2002.1039292","title":"Integrating cognitive support with CASE-tools for design recovery","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Diagrammatic reasoning; Component (thermodynamics); Task (project management); Process (computing); Software engineering; Simple (philosophy); Engineering design process; Human–computer interaction; Systems engineering; Programming language; Engineering","score_opus":0.05861207790964929,"score_gpt":0.2918170628448055,"score_spread":0.23320498493515623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158046712","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023325475,0.00034495845,0.92881256,0.0028427825,0.00010723813,0.0003658543,0.0001441223,0.005555501,0.03850149],"genre_scores_gemma":[0.14043605,0.00042729685,0.85303706,0.00024528333,0.000046081877,0.00030418072,0.00038457286,0.0003093171,0.0048101144],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99412835,0.0026110138,0.00058530434,0.0005438258,0.0017986877,0.00033271522],"domain_scores_gemma":[0.9620969,0.025644759,0.001888966,0.007970436,0.001685699,0.00071314693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00943362,0.001429003,0.0005965391,0.004944935,0.0014175234,0.008489998,0.0056981756,0.002605307,0.008368847],"category_scores_gemma":[0.034457896,0.00095731445,0.0012538095,0.0018839936,0.0036456939,0.008973263,0.008443998,0.0026173138,0.00194058],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023079288,0.0016707705,0.0028304048,0.000826917,0.00008570041,0.0032387248,0.01379076,0.02368353,0.0090693915,0.31866604,0.011936493,0.6139705],"study_design_scores_gemma":[0.000489948,0.00029070527,0.0019310137,0.0013472365,0.00021709758,0.0036423027,0.0044142893,0.2047665,0.018285839,0.44625923,0.31807107,0.00028473054],"about_ca_topic_score_codex":0.0020649051,"about_ca_topic_score_gemma":0.0037590833,"teacher_disagreement_score":0.00943362,"about_ca_system_score_codex":0.0016902576,"about_ca_system_score_gemma":0.0023747606,"threshold_uncertainty_score":0.04989028},"labels":[],"label_agreement":null},{"id":"W2158133897","doi":"10.1109/icsm.2008.4658082","title":"Duplicate bug reports considered harmful &amp;#x2026; really?","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":281,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Merge (version control); Computer science; Software bug; Open source; Data science; Security bug; World Wide Web; Information retrieval; Computer security; Software; Programming language; Information security","score_opus":0.05022993915979513,"score_gpt":0.2794636414728156,"score_spread":0.22923370231302048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158133897","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94187576,0.013247887,0.022594417,0.0066436315,0.0006862018,0.00027242693,0.0018177676,0.0015078994,0.011353958],"genre_scores_gemma":[0.97985566,0.002856373,0.011417253,0.0013764567,0.00035940885,0.00009480533,0.0011571319,0.00035062278,0.0025322286],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.97099864,0.008771855,0.0037586072,0.0029744175,0.012237974,0.0012585063],"domain_scores_gemma":[0.7992225,0.08969434,0.06695585,0.008989495,0.031611804,0.003525983],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013006685,0.00092924945,0.0012546297,0.0072543626,0.0015058009,0.0029127419,0.0010079402,0.001600331,0.0017034092],"category_scores_gemma":[0.1254446,0.000596803,0.0008654766,0.007892414,0.001300753,0.004209591,0.0023610697,0.0009811509,0.0006613274],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053230237,0.00016258864,0.5640739,0.0018228672,0.00044385233,0.0020506463,0.007727159,0.0010891982,0.004413952,0.0024602807,0.016107932,0.3991153],"study_design_scores_gemma":[0.00009736579,0.0009202894,0.8298226,0.002153101,0.0017306047,0.016815243,0.019714523,0.007604342,0.011811759,0.012864793,0.096023776,0.00044163366],"about_ca_topic_score_codex":0.002954076,"about_ca_topic_score_gemma":0.004386102,"teacher_disagreement_score":0.9869933,"about_ca_system_score_codex":0.0013842302,"about_ca_system_score_gemma":0.0016577417,"threshold_uncertainty_score":0.06878674},"labels":[],"label_agreement":null},{"id":"W2158236248","doi":"10.1109/apsec.2005.100","title":"Supporting predictive change impact analysis: a control call graph based technique","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Change impact analysis; Computer science; Call graph; Java; Regression testing; Control flow; Control flow graph; Program analysis; Software; Graph; Set (abstract data type); Software maintenance; Control (management); Static analysis; Data mining; Software engineering; Software system; Artificial intelligence; Theoretical computer science; Programming language; Software construction","score_opus":0.01523830419380241,"score_gpt":0.31120184765313286,"score_spread":0.29596354345933046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158236248","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01843412,0.000106287756,0.9609428,0.00020009436,0.00003464379,0.00012161316,0.00036405624,0.018174136,0.0016223363],"genre_scores_gemma":[0.5576005,0.00020403821,0.43808183,0.00017205269,0.000087556444,0.00023034065,0.0009898541,0.0010436483,0.001590129],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987602,0.00017948233,0.00005051073,0.00024055487,0.00068593543,0.000083391606],"domain_scores_gemma":[0.9918305,0.0049251067,0.00083449436,0.0015294387,0.0007809836,0.00009952015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009927446,0.0014981605,0.0008428732,0.0033792772,0.00050307706,0.00084512105,0.0015696707,0.00117149,0.0028532166],"category_scores_gemma":[0.009872632,0.00056331226,0.0009890703,0.0016720553,0.00067370303,0.002389103,0.00083783333,0.0014682051,0.00089335965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036413386,0.0004902726,0.013296458,0.00039282208,0.00023116176,0.000719262,0.00042360486,0.29583943,0.03562822,0.018089054,0.0072365166,0.62728906],"study_design_scores_gemma":[0.000019764395,0.00009116001,0.0017400956,0.000031972533,0.00006161522,0.00021066665,0.000020779187,0.9707118,0.014117967,0.010251233,0.0026951025,0.00004788947],"about_ca_topic_score_codex":0.0054049823,"about_ca_topic_score_gemma":0.0066497405,"teacher_disagreement_score":0.0054049823,"about_ca_system_score_codex":0.000505301,"about_ca_system_score_gemma":0.001138936,"threshold_uncertainty_score":0.0107470155},"labels":[],"label_agreement":null},{"id":"W2158324008","doi":"10.1109/csmr.2011.17","title":"Factbase and Decomposition Generation","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Cluster analysis; Program comprehension; Implementation; Software maintenance; Software; Data mining; Decomposition; Software metric; Measure (data warehouse); Software engineering; Software system; Machine learning; Software construction; Programming language","score_opus":0.05197087965235644,"score_gpt":0.2786029247926169,"score_spread":0.22663204514026045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158324008","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04531832,0.00039381874,0.9279794,0.0005131285,0.00019713334,0.0003969344,0.0025642165,0.011691861,0.0109451385],"genre_scores_gemma":[0.17301391,0.00022545783,0.81303585,0.00015128557,0.00005256958,0.00038415985,0.0065283068,0.0018791226,0.0047292393],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957337,0.0014093053,0.00023780191,0.00064259686,0.0017853864,0.00019134488],"domain_scores_gemma":[0.9758537,0.012427421,0.001326874,0.0069263084,0.003049964,0.0004158209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005111321,0.0007755015,0.0005701858,0.0032939734,0.0009406382,0.0030437314,0.0013957123,0.0010282128,0.008401462],"category_scores_gemma":[0.044154566,0.0005938069,0.0010722704,0.0027065484,0.0010051099,0.004132257,0.0021659127,0.0014689092,0.0020159197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065294857,0.000347762,0.009933772,0.0006785587,0.00014944866,0.00047513397,0.0018791375,0.04791081,0.012175055,0.25438967,0.036291327,0.63511634],"study_design_scores_gemma":[0.00016734614,0.0002937246,0.0067781135,0.00026887344,0.00014701777,0.0009780321,0.0005654993,0.515653,0.048978943,0.26959383,0.15644877,0.00012684267],"about_ca_topic_score_codex":0.0023403096,"about_ca_topic_score_gemma":0.0029322458,"teacher_disagreement_score":0.008401462,"about_ca_system_score_codex":0.0010329344,"about_ca_system_score_gemma":0.0016820519,"threshold_uncertainty_score":0.028105676},"labels":[],"label_agreement":null},{"id":"W2158324307","doi":"10.1109/icsm.2007.4362614","title":"Mining the Lexicon Used by Programmers during Sofware Evolution","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Identifier; Lexicon; Computer science; Program comprehension; Documentation; Process (computing); Task (project management); Comprehension; Natural language processing; Artificial intelligence; Information retrieval; Programming language; Software; Software system","score_opus":0.015670194518798263,"score_gpt":0.26420079313084227,"score_spread":0.24853059861204402,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158324307","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9815477,0.00030288336,0.014667094,0.00015001478,0.000012513942,0.00008492253,0.0009190055,0.0005260862,0.0017897706],"genre_scores_gemma":[0.94929814,0.00028818144,0.04368864,0.000057522946,0.000010802171,0.00017386253,0.0045098853,0.00031172534,0.00166123],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99680173,0.000996361,0.0004222626,0.00061133254,0.0009375159,0.00023068697],"domain_scores_gemma":[0.9747274,0.014622554,0.0040377537,0.0019029374,0.004268892,0.0004404471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028656228,0.00040288808,0.000561004,0.00731336,0.00093907025,0.0029156308,0.00076462456,0.0007729569,0.00055459567],"category_scores_gemma":[0.033265796,0.0004668029,0.0005125348,0.005814171,0.001013685,0.003565874,0.0013717134,0.00075379305,0.0004178968],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041531693,0.0002018318,0.52007467,0.000664119,0.000104623614,0.0020789707,0.027972013,0.0042668693,0.03769985,0.0031991706,0.0034237746,0.3998988],"study_design_scores_gemma":[0.0000843051,0.00065214717,0.7249122,0.0003447791,0.0004242587,0.004682461,0.026806273,0.13117523,0.053005278,0.012702284,0.044853527,0.00035723406],"about_ca_topic_score_codex":0.008043495,"about_ca_topic_score_gemma":0.011225089,"teacher_disagreement_score":0.008043495,"about_ca_system_score_codex":0.0015270254,"about_ca_system_score_gemma":0.0015846331,"threshold_uncertainty_score":0.015993357},"labels":[],"label_agreement":null},{"id":"W2158344808","doi":"10.1109/ictai.2009.110","title":"Using Concepts Analysis for Mining Functional Features from Legacy Code","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Implementation; Code (set theory); Legacy system; Programming language; Legacy code; Inheritance (genetic algorithm); Software engineering; Source code; Software","score_opus":0.06712511127362189,"score_gpt":0.34408607001268127,"score_spread":0.27696095873905935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158344808","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26433924,0.0010538425,0.7247508,0.0003448855,0.00003215614,0.0010081449,0.003003654,0.002505541,0.0029616898],"genre_scores_gemma":[0.26815018,0.000283787,0.7257369,0.00005664336,0.00002345232,0.0007011589,0.0042164708,0.00015862421,0.0006727409],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982033,0.00030098655,0.00022475854,0.00045686122,0.0007070123,0.000106982945],"domain_scores_gemma":[0.9886975,0.0067704106,0.001492213,0.00093237613,0.0018349459,0.00027250766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024930409,0.0011058078,0.0005799814,0.015904548,0.0010728402,0.0015417859,0.0014396721,0.00090919674,0.0011314513],"category_scores_gemma":[0.015046756,0.0004294178,0.0011983127,0.0058683227,0.0012024788,0.002768999,0.0014865864,0.0009007495,0.00040819132],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000353511,0.0005665777,0.101853,0.0019225174,0.00028509554,0.0028763288,0.0074260943,0.02284711,0.03089773,0.029937623,0.004185904,0.79684854],"study_design_scores_gemma":[0.00025467528,0.001059758,0.120565906,0.0008856167,0.0006734096,0.0061085504,0.007287098,0.61570424,0.07088182,0.12109918,0.05503711,0.00044265448],"about_ca_topic_score_codex":0.0058321645,"about_ca_topic_score_gemma":0.0057773916,"teacher_disagreement_score":0.015904548,"about_ca_system_score_codex":0.0010619517,"about_ca_system_score_gemma":0.0024803658,"threshold_uncertainty_score":0.013184607},"labels":[],"label_agreement":null},{"id":"W2158423705","doi":"10.1109/icsm.2007.4362671","title":"Improving Predictive Models of Software Quality Using an Evolutionary Computational Approach","year":2007,"lang":"en","type":"article","venue":"Proceedings/Proceedings - Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Classifier (UML); Machine learning; Artificial intelligence; Feature selection; Software quality; Data mining; Predictive modelling; Software; Software development","score_opus":0.06543109641917878,"score_gpt":0.30256784052823943,"score_spread":0.23713674410906066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158423705","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15906328,0.00042445114,0.8349782,0.0011350672,0.000050845876,0.00013915126,0.00013429209,0.0007506823,0.0033240132],"genre_scores_gemma":[0.85476506,0.00031145857,0.14286624,0.00022544015,0.00006892231,0.0002478435,0.00027464348,0.000064132044,0.0011762773],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989128,0.00046876734,0.000056799523,0.00019764065,0.00027907256,0.00008485288],"domain_scores_gemma":[0.98629653,0.011371207,0.0006754281,0.0004840475,0.0010626726,0.000110024266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041422034,0.0011271815,0.0012343968,0.002509256,0.00058087736,0.0017756561,0.002095794,0.0017131608,0.0013444743],"category_scores_gemma":[0.024109326,0.00079678616,0.0009081128,0.0014460573,0.00087496446,0.0022972408,0.0009346914,0.0019132257,0.00026763204],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017381752,0.00004762897,0.0017409408,0.000012522317,0.000030929074,0.000020360356,0.00002533585,0.97975254,0.00017952904,0.0019224182,0.00015683177,0.016093541],"study_design_scores_gemma":[0.000001909835,0.000005642294,0.00009333135,0.000002507495,0.0000033494077,0.0000029195699,0.0000022044019,0.998575,0.000036534133,0.0012489394,0.00002606752,0.0000015321228],"about_ca_topic_score_codex":0.009505419,"about_ca_topic_score_gemma":0.0070322403,"teacher_disagreement_score":0.009505419,"about_ca_system_score_codex":0.0014219324,"about_ca_system_score_gemma":0.0010443486,"threshold_uncertainty_score":0.021906376},"labels":[],"label_agreement":null},{"id":"W2158439356","doi":"10.1109/icsm.2015.7332459","title":"Evaluating clone detection tools with BigCloneBench","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":164,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Code refactoring; Benchmark (surveying); Java; Computer science; Precision and recall; Software; Similarity (geometry); Source code; Data mining; Software engineering; Artificial intelligence; Programming language; Biology; Genetics; Gene","score_opus":0.13330618489116314,"score_gpt":0.35197360513834863,"score_spread":0.2186674202471855,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158439356","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.71203965,0.006534145,0.16592477,0.00084848434,0.00039035807,0.00079100375,0.010513111,0.09692804,0.0060304273],"genre_scores_gemma":[0.6505093,0.0011059536,0.30149186,0.0005907708,0.00012255649,0.00076523854,0.03976301,0.0035947885,0.002056504],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9723388,0.005812797,0.004488592,0.0038280096,0.01274669,0.0007851083],"domain_scores_gemma":[0.83516693,0.105492316,0.013629804,0.022239199,0.021116968,0.0023547325],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.019021701,0.0025081555,0.0013676476,0.015752325,0.0009415227,0.0044296817,0.004382661,0.003108435,0.0009963475],"category_scores_gemma":[0.11143408,0.00093772466,0.0017376265,0.008120441,0.0014621504,0.006067866,0.003953103,0.0014814177,0.0008368986],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018456577,0.0018873157,0.31265697,0.0048629274,0.002212756,0.0017421447,0.0049032043,0.056434188,0.040488727,0.007222842,0.052980643,0.5127627],"study_design_scores_gemma":[0.0007870376,0.0035421876,0.109237425,0.0011107421,0.0009461552,0.004556049,0.0025711746,0.66462696,0.14202972,0.009375382,0.060449384,0.0007677321],"about_ca_topic_score_codex":0.0049532237,"about_ca_topic_score_gemma":0.0057929256,"teacher_disagreement_score":0.9809783,"about_ca_system_score_codex":0.0016168755,"about_ca_system_score_gemma":0.0020837758,"threshold_uncertainty_score":0.10059762},"labels":[],"label_agreement":null},{"id":"W2158485104","doi":"10.1109/cmpsac.2002.1045082","title":"Object identification in legacy code as a grouping problem","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Legacy system; Computer science; Automation; Software engineering; Task (project management); Legacy code; Software maintenance; Process (computing); Identification (biology); Entropy (arrow of time); Software; Source code; Code (set theory); Programming language; Software system; Distributed computing; Systems engineering; Engineering","score_opus":0.019897031564289586,"score_gpt":0.28300597822971174,"score_spread":0.26310894666542217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158485104","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25376418,0.00038732606,0.7383033,0.0006370355,0.000068400695,0.00020812538,0.00012883569,0.0027724789,0.0037303122],"genre_scores_gemma":[0.38934755,0.00023066245,0.60381037,0.00012343116,0.000040371928,0.00010290453,0.0004864815,0.00033572505,0.0055225133],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99819225,0.00051599153,0.00012761739,0.00041890645,0.0005870204,0.00015832442],"domain_scores_gemma":[0.9913482,0.0043445053,0.0014107771,0.001453826,0.0011115756,0.00033112854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023729028,0.00069872325,0.0009517841,0.002368455,0.0022615155,0.0024093702,0.0017849783,0.0022777878,0.0018161573],"category_scores_gemma":[0.0100904675,0.0005206669,0.0006596853,0.0027232831,0.001536327,0.0034792887,0.0022244882,0.00094209745,0.0009663565],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059068156,0.00040579258,0.021606836,0.0003121249,0.000054447984,0.0011141597,0.0039188433,0.10642682,0.024707565,0.0440401,0.009385969,0.7874366],"study_design_scores_gemma":[0.00012852447,0.0002961226,0.011709721,0.000106826046,0.00009306876,0.0018078253,0.0029477365,0.8268635,0.036947276,0.09472865,0.024253065,0.00011764164],"about_ca_topic_score_codex":0.003418304,"about_ca_topic_score_gemma":0.0028161982,"teacher_disagreement_score":0.003418304,"about_ca_system_score_codex":0.0007869011,"about_ca_system_score_gemma":0.00097373204,"threshold_uncertainty_score":0.012549281},"labels":[],"label_agreement":null},{"id":"W2158694364","doi":"10.1109/wcre.2012.27","title":"Reverse Engineering iOS Mobile Applications","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Reverse engineering; Mobile device; Popularity; Software; User interface; Interface (matter); Cover (algebra); State (computer science); Mobile computing; Human–computer interaction; Software engineering; World Wide Web; Operating system; Programming language; Engineering","score_opus":0.011419722251319268,"score_gpt":0.25277387335421814,"score_spread":0.24135415110289887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158694364","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16422896,0.00036949888,0.79381716,0.000340029,0.00010008256,0.0005407936,0.0004172919,0.03271722,0.0074689505],"genre_scores_gemma":[0.5052726,0.00039768184,0.4777556,0.00027210207,0.000026512173,0.00029787762,0.0012159946,0.003789987,0.010971644],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99859875,0.00029845245,0.00010957131,0.0002618466,0.00057874946,0.0001526732],"domain_scores_gemma":[0.9935382,0.0024978293,0.00069089886,0.002084409,0.0011120675,0.00007671246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012875746,0.0010356985,0.0004014979,0.0011230862,0.00039883994,0.0010563361,0.0011049239,0.00089588016,0.0021606404],"category_scores_gemma":[0.008644006,0.00071178604,0.00092605397,0.00039922947,0.00088189554,0.0016401183,0.0014877489,0.001151348,0.0009388608],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005583168,0.0003637955,0.024389533,0.0012264728,0.00015613237,0.0022614864,0.0036017133,0.05705272,0.25975856,0.029185267,0.008145179,0.61330086],"study_design_scores_gemma":[0.000069782815,0.00054251036,0.005990825,0.0002514733,0.00020206679,0.0022013932,0.0006831934,0.53737384,0.37228072,0.020290973,0.060008824,0.00010446909],"about_ca_topic_score_codex":0.002071641,"about_ca_topic_score_gemma":0.0038069657,"teacher_disagreement_score":0.0021606404,"about_ca_system_score_codex":0.00055869424,"about_ca_system_score_gemma":0.0011603318,"threshold_uncertainty_score":0.0072280765},"labels":[],"label_agreement":null},{"id":"W2158735796","doi":"10.1109/icsm.2006.35","title":"Managing Concern Interfaces","year":2006,"lang":"en","type":"article","venue":"Proceedings/Proceedings - Conference on Software Maintenance","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Interface (matter); User interface; Application programming interface; Relation (database); Information hiding; Separation of concerns; Software; Software engineering; Human–computer interaction; Interface description language; Programming language; Database; Artificial intelligence; Operating system","score_opus":0.027339568136747355,"score_gpt":0.2590262910534106,"score_spread":0.23168672291666323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158735796","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01784492,0.00062609214,0.95851237,0.000798269,0.00014723302,0.00045275505,0.00016773952,0.009010213,0.01244032],"genre_scores_gemma":[0.23360424,0.0008391109,0.7370075,0.0010217751,0.0002371758,0.0007786506,0.0010897192,0.0045993803,0.020822559],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9903128,0.0026482749,0.0013249036,0.0013086409,0.003467307,0.0009381346],"domain_scores_gemma":[0.98175365,0.0053214193,0.001700908,0.0076261326,0.002863861,0.00073409843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009885684,0.0016145125,0.00089033064,0.0029579015,0.0018595709,0.0074328952,0.00433337,0.002558043,0.004836446],"category_scores_gemma":[0.034673583,0.0014005951,0.001601254,0.0013631606,0.0015029015,0.009800632,0.010800109,0.003766521,0.002236166],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003336049,0.000339541,0.012528869,0.0008686058,0.00025719978,0.0013147717,0.011786898,0.004961172,0.03750909,0.3673244,0.021119175,0.54165673],"study_design_scores_gemma":[0.00018911008,0.00033867455,0.0045859255,0.0006387306,0.00072503986,0.0032480601,0.0018787165,0.07427282,0.08639428,0.28895783,0.5384793,0.0002915442],"about_ca_topic_score_codex":0.0013694703,"about_ca_topic_score_gemma":0.0010982503,"teacher_disagreement_score":0.009885684,"about_ca_system_score_codex":0.0013135469,"about_ca_system_score_gemma":0.0020665783,"threshold_uncertainty_score":0.05228114},"labels":[],"label_agreement":null},{"id":"W2158744032","doi":"10.1109/icse.2009.5070510","title":"Predicting faults using the complexity of code changes","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":697,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Code (set theory); Source code; Process (computing); Software; Product metric; Product (mathematics); Software bug; Reliability engineering; Programming language; Mathematics; Engineering; Set (abstract data type)","score_opus":0.11646171279775151,"score_gpt":0.32917074110240396,"score_spread":0.21270902830465244,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158744032","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9815154,0.00026533558,0.015914293,0.0001376711,0.000017234664,0.000059337035,0.00090120843,0.00024215478,0.00094734767],"genre_scores_gemma":[0.99512035,0.00007857835,0.0038108875,0.0000074424124,0.000018048595,0.000020992979,0.00081505737,0.000014625595,0.00011405584],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.996775,0.0005993901,0.00041173055,0.0005881465,0.001428903,0.00019688898],"domain_scores_gemma":[0.89459544,0.06315182,0.027912056,0.0042077233,0.008060781,0.0020721194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022963907,0.0007855248,0.0005810997,0.010823133,0.0003435354,0.0014924415,0.000522892,0.0010394161,0.00074596464],"category_scores_gemma":[0.052231885,0.00033552633,0.0006197834,0.0042958236,0.0006028926,0.0033285727,0.0009570828,0.0008777864,0.00032862186],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011553211,0.00010325924,0.9197813,0.00006002667,0.00012690619,0.00010655767,0.00015777984,0.03881903,0.001613876,0.0003222485,0.0003150282,0.03847841],"study_design_scores_gemma":[0.00001741084,0.00043593053,0.66421926,0.000028210337,0.0000652567,0.00035782816,0.00018344155,0.32835913,0.0032935652,0.0021847265,0.00079369283,0.0000615042],"about_ca_topic_score_codex":0.0036455642,"about_ca_topic_score_gemma":0.0051702834,"teacher_disagreement_score":0.010823133,"about_ca_system_score_codex":0.00087458774,"about_ca_system_score_gemma":0.0005297728,"threshold_uncertainty_score":0.012144625},"labels":[],"label_agreement":null},{"id":"W2158925837","doi":"10.1109/icsm.2007.4362664","title":"Software Artefact Traceability: the Never-Ending Challenge","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Traceability; Computer science; Software development; Software engineering; TRACE (psycholinguistics); Software; Software maintenance; Software system; Session (web analytics); Task (project management); Software construction; Requirements traceability; Verification and validation; Package development process; Systems engineering; Engineering; World Wide Web; Operating system","score_opus":0.028532969609841675,"score_gpt":0.2811823337077057,"score_spread":0.252649364097864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158925837","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019319177,0.03729942,0.7373578,0.18201423,0.0029523235,0.00013809143,0.0001811615,0.0015090675,0.019228645],"genre_scores_gemma":[0.4178197,0.06658772,0.46170762,0.0163747,0.009976688,0.00041170247,0.0006577331,0.0018198275,0.024644254],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9719052,0.011122371,0.001768597,0.0029798846,0.011258115,0.00096583535],"domain_scores_gemma":[0.8193299,0.12167055,0.0055160066,0.02852989,0.020967942,0.0039856243],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031279534,0.0013837887,0.0029036961,0.0030804658,0.0051055504,0.014420142,0.005397231,0.009740344,0.004351886],"category_scores_gemma":[0.0999513,0.0011275315,0.0014934103,0.003066029,0.015461636,0.045520928,0.009503768,0.011845103,0.0022355788],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013695208,0.00020308032,0.003608033,0.0015231606,0.00013678124,0.001095144,0.004928037,0.0065371306,0.002156264,0.45095247,0.026649183,0.5020737],"study_design_scores_gemma":[0.00001981185,0.0000843783,0.00062809454,0.0006569274,0.000033122902,0.0014416865,0.0022183186,0.009277582,0.0015660416,0.9004427,0.08353781,0.00009351448],"about_ca_topic_score_codex":0.0029441027,"about_ca_topic_score_gemma":0.0020302874,"teacher_disagreement_score":0.031279534,"about_ca_system_score_codex":0.0029619301,"about_ca_system_score_gemma":0.0065935343,"threshold_uncertainty_score":0.16542399},"labels":[],"label_agreement":null},{"id":"W2158949715","doi":"10.1002/0471028959.sof282","title":"Resource Estimation in Software Engineering","year":2002,"lang":"en","type":"other","venue":"Encyclopedia of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":146,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Estimation; Computer science; Resource (disambiguation); Sizing; Software; Cost estimate; Management science; Operations research; Data science; Software engineering; Industrial engineering; Data mining; Systems engineering; Engineering","score_opus":0.007205205658071253,"score_gpt":0.21413443820999986,"score_spread":0.2069292325519286,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158949715","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015476081,0.012702281,0.94131005,0.0018224152,0.00022837702,0.00033744035,0.0003130971,0.00068202114,0.027128236],"genre_scores_gemma":[0.4848668,0.0089155985,0.497962,0.00044050996,0.0003166683,0.00081192853,0.00078960933,0.0003838845,0.0055129854],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.95064926,0.031829398,0.002323214,0.002449982,0.011976177,0.0007720004],"domain_scores_gemma":[0.9188402,0.057852242,0.0055891047,0.004961076,0.012177014,0.0005802648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026344918,0.0017983057,0.0016510237,0.010543195,0.0010414708,0.006162505,0.0020936492,0.0014752432,0.00403723],"category_scores_gemma":[0.109489866,0.00062670076,0.0013219751,0.010243007,0.0026513555,0.007329853,0.0039476766,0.0017385597,0.00071941846],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008867191,0.00011068684,0.008297511,0.0016081963,0.0002117128,0.00008012214,0.0008220794,0.13241884,0.0007192779,0.22773881,0.006028195,0.62187594],"study_design_scores_gemma":[0.00003414252,0.00031699394,0.011210782,0.0034250694,0.00021836201,0.00016694135,0.0012835489,0.47974962,0.004724563,0.4397991,0.058876462,0.00019451474],"about_ca_topic_score_codex":0.008311808,"about_ca_topic_score_gemma":0.0045291586,"teacher_disagreement_score":0.026344918,"about_ca_system_score_codex":0.005064934,"about_ca_system_score_gemma":0.003982605,"threshold_uncertainty_score":0.13932687},"labels":[],"label_agreement":null},{"id":"W2159294330","doi":"10.1109/msr.2010.5463345","title":"Should I contribute to this discussion?","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Python (programming language); Computer science; World Wide Web; Open source; Open-source software development; Java; Data science; Software engineering; Software; Programming language","score_opus":0.019469058295621813,"score_gpt":0.2984373018064,"score_spread":0.2789682435107782,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2159294330","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057127014,0.006083742,0.03161171,0.5345648,0.034237746,0.0005078478,0.001127947,0.0018907555,0.33284843],"genre_scores_gemma":[0.516173,0.010928094,0.02691858,0.10915198,0.027368143,0.00075530336,0.0014905629,0.0015014412,0.30571285],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9929789,0.0039846054,0.00024794976,0.00074285,0.0013337259,0.0007120059],"domain_scores_gemma":[0.9467985,0.020210855,0.006244711,0.0033909287,0.008054155,0.015300873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010612879,0.00062360615,0.00064170803,0.0011961181,0.0035505076,0.0062045665,0.0013968319,0.0046488526,0.043021303],"category_scores_gemma":[0.07970482,0.0004102566,0.00077038864,0.0010493951,0.0023454574,0.008302903,0.0032831002,0.0056402716,0.034426812],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038424216,0.00036615581,0.032409128,0.00046189642,0.00008047542,0.0013257348,0.035568446,0.00025350854,0.0013511239,0.044908345,0.6707125,0.21217845],"study_design_scores_gemma":[0.000058409514,0.0001877764,0.0047195014,0.0006443194,0.000081531936,0.001345279,0.017460844,0.00089603884,0.00060954835,0.028551806,0.9453652,0.000079757265],"about_ca_topic_score_codex":0.001082007,"about_ca_topic_score_gemma":0.0012436197,"teacher_disagreement_score":0.043021303,"about_ca_system_score_codex":0.001530321,"about_ca_system_score_gemma":0.0033813578,"threshold_uncertainty_score":0.14392054},"labels":[],"label_agreement":null},{"id":"W2159566003","doi":"10.1002/smr.245","title":"Field studies using functional size measurement in building estimation models for software maintenance","year":2002,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software; Computer science; Field (mathematics); Estimation; Functional requirement; Software development; Software metric; Function point; Software maintenance; Software engineering; Reliability engineering; Systems engineering; Industrial engineering; Engineering; Software construction; Mathematics","score_opus":0.19941479076137206,"score_gpt":0.3802488116615875,"score_spread":0.18083402090021547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2159566003","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8119316,0.00024773128,0.18590766,0.00014010466,0.000019285431,0.00023798017,0.00016035345,0.00016759954,0.0011877093],"genre_scores_gemma":[0.9366561,0.00011611359,0.062512934,0.000025338655,0.000017806358,0.00018511138,0.00029356562,0.000017111834,0.00017591935],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.983307,0.013993779,0.00053950463,0.0011538705,0.0008602321,0.0001454864],"domain_scores_gemma":[0.6731148,0.30119005,0.009927453,0.009132167,0.0060298145,0.0006057607],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.029104726,0.00089931017,0.0006037137,0.0029750722,0.00040907462,0.000912213,0.0015125017,0.0008845958,0.0009601017],"category_scores_gemma":[0.112366155,0.0005848938,0.00082191237,0.0016978743,0.001030273,0.0019589383,0.0010236118,0.0010866108,0.00017261304],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010879098,0.0025075432,0.46667916,0.0005273647,0.0008932841,0.0001732913,0.0032731364,0.27643403,0.0036219505,0.011868674,0.0016550308,0.23127867],"study_design_scores_gemma":[0.00022677406,0.0026436655,0.10307169,0.0001479628,0.00019093469,0.000229533,0.0014236271,0.87368286,0.004821112,0.011034541,0.0023891625,0.00013819747],"about_ca_topic_score_codex":0.0041346396,"about_ca_topic_score_gemma":0.003487833,"teacher_disagreement_score":0.9708953,"about_ca_system_score_codex":0.0011378531,"about_ca_system_score_gemma":0.0005283862,"threshold_uncertainty_score":0.15392232},"labels":[],"label_agreement":null},{"id":"W2159903682","doi":"10.1145/337180.337477","title":"An evaluation of the paired comparisons method for software sizing","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Research Canada; Ericsson (Canada)","funders":"","keywords":"Sizing; Computer science; Software; Reliability engineering; Engineering; Programming language; Chemistry","score_opus":0.08624099027380454,"score_gpt":0.37930923294804414,"score_spread":0.2930682426742396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2159903682","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039817292,0.0005877415,0.9513083,0.00010549632,0.0005170414,0.00055506907,0.00033536844,0.0012337293,0.005539913],"genre_scores_gemma":[0.40507826,0.00026539539,0.5895556,0.00013753104,0.0002298804,0.0014673084,0.00051088096,0.0006579693,0.0020972071],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9132049,0.062028475,0.0018869872,0.005598607,0.016690379,0.0005906451],"domain_scores_gemma":[0.64894557,0.3085393,0.005781186,0.015205954,0.020760532,0.00076741516],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.051935434,0.0013485696,0.0013227585,0.002823339,0.0011518406,0.0016866191,0.0021241817,0.0012593522,0.0056828237],"category_scores_gemma":[0.20752978,0.0006108746,0.0012779624,0.0021005785,0.0021553268,0.0018047187,0.0016652017,0.0017470596,0.0012336563],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.015672103,0.00093316246,0.02632245,0.001246009,0.0019845932,0.00043909592,0.0016828086,0.05349815,0.016715145,0.02896939,0.008376217,0.84416085],"study_design_scores_gemma":[0.0021372687,0.029933708,0.054086827,0.00039961797,0.0013922979,0.0033189058,0.0014043445,0.717279,0.109617524,0.051860638,0.027836682,0.00073319906],"about_ca_topic_score_codex":0.0007879696,"about_ca_topic_score_gemma":0.0009338484,"teacher_disagreement_score":0.94806457,"about_ca_system_score_codex":0.0009306224,"about_ca_system_score_gemma":0.00087519205,"threshold_uncertainty_score":0.2746641},"labels":[],"label_agreement":null},{"id":"W2160184302","doi":"10.1016/j.infsof.2009.10.003","title":"An effort prediction framework for software defect correction","year":2009,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software bug; Computer science; Software; Reliability engineering; Software engineering; Programming language; Engineering","score_opus":0.007486731375790292,"score_gpt":0.2595437837150833,"score_spread":0.252057052339293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160184302","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011412954,0.0002473963,0.98461443,0.00016504634,0.000036435416,0.000056153975,0.00022806478,0.002213984,0.0010254917],"genre_scores_gemma":[0.5438171,0.00039393766,0.4510108,0.00012093849,0.0001114736,0.00018595604,0.0007471917,0.0001644845,0.0034481063],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990332,0.00016782031,0.000060984865,0.00021880318,0.00039843237,0.000120665536],"domain_scores_gemma":[0.9981025,0.0007940308,0.000184709,0.00019140924,0.0006417306,0.00008579565],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017387534,0.00093999173,0.001307337,0.002163936,0.0005057867,0.0012518407,0.0019275895,0.000959981,0.0017235334],"category_scores_gemma":[0.004563187,0.00037164646,0.0008646935,0.0012694031,0.00043446748,0.001874702,0.0009562324,0.001039422,0.0004929183],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017720734,0.00042929928,0.007313744,0.00017094742,0.00019185065,0.00029619064,0.0001738367,0.49364695,0.0042106155,0.04934167,0.008681803,0.4353658],"study_design_scores_gemma":[0.0000059407766,0.00002260289,0.0004099932,0.000008503839,0.00001953237,0.00002385668,0.000009150588,0.98938876,0.00051121454,0.009020968,0.0005697886,0.000009713922],"about_ca_topic_score_codex":0.01583106,"about_ca_topic_score_gemma":0.016563324,"teacher_disagreement_score":0.01583106,"about_ca_system_score_codex":0.0008726407,"about_ca_system_score_gemma":0.0017769836,"threshold_uncertainty_score":0.03147781},"labels":[],"label_agreement":null},{"id":"W2160381050","doi":"10.1145/337180.337823","title":"Beg, borrow, or steal (workshop session)","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of New Brunswick; University of Victoria; National Research Council Canada","funders":"","keywords":"Session (web analytics); Computer science; Computer security; World Wide Web","score_opus":0.0243000999995121,"score_gpt":0.2951911096943852,"score_spread":0.2708910096948731,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160381050","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013451732,0.01927425,0.017862737,0.09579694,0.38304922,0.0029456902,0.006492888,0.005047914,0.4560787],"genre_scores_gemma":[0.023576932,0.0064785792,0.004925662,0.028008096,0.030582424,0.0015632425,0.003575797,0.001229867,0.90005946],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999466,0.0001182809,0.000024518577,0.00012367973,0.000117227995,0.00015032633],"domain_scores_gemma":[0.9977623,0.00025435243,0.000089034154,0.000088977795,0.0004605196,0.0013447283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002730498,0.0018243307,0.0009568038,0.00066671567,0.002359288,0.0059440294,0.0014967838,0.0051249177,0.27438253],"category_scores_gemma":[0.00401451,0.0005818296,0.0014737487,0.0005006946,0.00057859777,0.0037599728,0.005192473,0.005385786,0.19861051],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002951413,0.000090375834,0.00008578771,0.00014669726,0.000010414994,0.00017382028,0.00019053138,0.000073479445,0.0008189903,0.0008981485,0.97448593,0.022730794],"study_design_scores_gemma":[0.000072308714,0.00011531893,0.00056902453,0.00019832354,0.00001326864,0.00009976034,0.00059905375,0.00010335485,0.00035757473,0.0013406533,0.99650156,0.000029852077],"about_ca_topic_score_codex":0.001155516,"about_ca_topic_score_gemma":0.0028690393,"teacher_disagreement_score":0.27438253,"about_ca_system_score_codex":0.000963584,"about_ca_system_score_gemma":0.0012191887,"threshold_uncertainty_score":0.91790104},"labels":[],"label_agreement":null},{"id":"W2160423649","doi":"10.1109/isese.2005.1541811","title":"Managing software change tasks: an exploratory study","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Software development; Software engineering; Java; Exploratory research; Context (archaeology); Software maintenance; Software; Software evolution; Software construction; Software system; Key (lock); Human–computer interaction; Programming language; Operating system","score_opus":0.05930237375680696,"score_gpt":0.30176481334598365,"score_spread":0.24246243958917668,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160423649","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9960259,0.000078422614,0.0021137972,0.00033214918,0.000006905518,0.00023439163,0.000038286486,0.00002967576,0.0011404102],"genre_scores_gemma":[0.99358314,0.00028293396,0.0038369836,0.00045581843,0.00002187701,0.00057280005,0.00007294987,0.00003200772,0.0011413958],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9916061,0.0056477566,0.0003778754,0.0006233274,0.0008328852,0.00091199076],"domain_scores_gemma":[0.9410898,0.04722522,0.0035948188,0.0018545482,0.0034070264,0.002828616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010987518,0.0010299074,0.0008092387,0.0018018571,0.007452488,0.0027334108,0.002423505,0.0030844985,0.0013623203],"category_scores_gemma":[0.045434672,0.0012414432,0.00037948095,0.0012360216,0.004475298,0.004149866,0.0035689282,0.00372456,0.0004451732],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021233816,0.0032794373,0.024169482,0.00033695527,0.00001895396,0.0024612593,0.94720036,0.00022364203,0.0058858837,0.00078497763,0.00079074845,0.014635987],"study_design_scores_gemma":[0.00014380999,0.0032296963,0.031033128,0.00025892243,0.000029614192,0.0021667425,0.9438373,0.001572096,0.0033386895,0.0011992616,0.01309129,0.000099348625],"about_ca_topic_score_codex":0.002753369,"about_ca_topic_score_gemma":0.005860094,"teacher_disagreement_score":0.010987518,"about_ca_system_score_codex":0.0014679737,"about_ca_system_score_gemma":0.002065065,"threshold_uncertainty_score":0.05810827},"labels":[],"label_agreement":null},{"id":"W2160506632","doi":"10.1109/wcre.2009.28","title":"An Exploratory Study of the Impact of Code Smells on Software Change-proneness","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":305,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code smell; Eclipse; Computer science; Software quality; Code (set theory); Relation (database); Software; Programming language; Software development; Data mining","score_opus":0.05944708631218296,"score_gpt":0.33589836101851983,"score_spread":0.27645127470633685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160506632","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99882656,0.00003157921,0.00058057776,0.000023895322,0.0000011735146,0.000017787394,0.00012935269,0.000021950491,0.0003670456],"genre_scores_gemma":[0.9990514,0.000016883783,0.00063073274,0.0000102094355,0.0000034142774,0.000014086535,0.00014427816,0.000008386783,0.00012063944],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9961423,0.0017069187,0.00029048848,0.00055462145,0.0010199425,0.00028573652],"domain_scores_gemma":[0.80808437,0.1500471,0.02844089,0.005412968,0.0052677332,0.0027469005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038666332,0.0004502087,0.00029744784,0.0019005385,0.00035151315,0.0007663702,0.00046511577,0.0005351357,0.0014481209],"category_scores_gemma":[0.042058114,0.00026760143,0.0005886688,0.0014762585,0.00060794555,0.0012418311,0.00069068203,0.00070788583,0.00024700788],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039191142,0.00031515246,0.98350567,0.00007134372,0.0001302518,0.00030751128,0.0012972056,0.0004662618,0.0054259896,0.00007760649,0.00010314089,0.007907857],"study_design_scores_gemma":[0.000006543805,0.0005242011,0.99672663,0.0000062129157,0.000023380311,0.00016203837,0.00033591222,0.0011144399,0.0009541978,0.000046283174,0.00009109071,0.000008981769],"about_ca_topic_score_codex":0.0011164083,"about_ca_topic_score_gemma":0.0019863076,"teacher_disagreement_score":0.0038666332,"about_ca_system_score_codex":0.00035628097,"about_ca_system_score_gemma":0.00038114676,"threshold_uncertainty_score":0.020448923},"labels":[],"label_agreement":null},{"id":"W2160517961","doi":"10.1145/1453101.1453146","title":"What makes a good bug report?","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":603,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Victoria","funders":"","keywords":"Eclipse; Computer science; Software bug; Software engineering; Software development; Quality (philosophy); Software; World Wide Web; Data science; Programming language","score_opus":0.028987268854475508,"score_gpt":0.27060477922221365,"score_spread":0.24161751036773815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160517961","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8581334,0.021842763,0.021699717,0.06644038,0.0031239756,0.0005090083,0.0007129864,0.0028398612,0.02469798],"genre_scores_gemma":[0.9775718,0.0041057607,0.01022052,0.0042851814,0.0014672744,0.000100346915,0.00040487756,0.00040648776,0.0014378708],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.94987863,0.017309265,0.00543294,0.0036084554,0.021520365,0.0022502884],"domain_scores_gemma":[0.5996835,0.23238751,0.095394775,0.009515469,0.05080236,0.0122164665],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030491859,0.0012354432,0.0013320899,0.008156376,0.0030379044,0.005222734,0.0015461271,0.0032007152,0.00261964],"category_scores_gemma":[0.26929995,0.0010494591,0.0009647356,0.0046819383,0.004977733,0.0117533365,0.0022544898,0.0020650122,0.001146942],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045132515,0.0007731222,0.60203505,0.0022514989,0.0003782658,0.0030822216,0.021871263,0.00053710956,0.0034737468,0.0049504503,0.040991742,0.31920415],"study_design_scores_gemma":[0.00016662474,0.0021080675,0.817842,0.0035273938,0.00090683764,0.015033302,0.058986664,0.0026636897,0.0041951523,0.018132964,0.07577916,0.00065822166],"about_ca_topic_score_codex":0.003016607,"about_ca_topic_score_gemma":0.004536215,"teacher_disagreement_score":0.030491859,"about_ca_system_score_codex":0.0022981307,"about_ca_system_score_gemma":0.0035298935,"threshold_uncertainty_score":0.16125828},"labels":[],"label_agreement":null},{"id":"W2160610963","doi":"10.1145/1718918.1718972","title":"Communication, collaboration, and bugs","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":147,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science","score_opus":0.007624884070102676,"score_gpt":0.2655405085967691,"score_spread":0.2579156245266664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160610963","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7946096,0.009828957,0.02524829,0.027224397,0.00032774734,0.00013117885,0.000104276354,0.0003324649,0.1421931],"genre_scores_gemma":[0.9937987,0.0014228747,0.002242973,0.00038427804,0.000083167164,0.00005565992,0.000026472786,0.000032544296,0.001953369],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.96590525,0.027308747,0.0010746258,0.001438177,0.0027975803,0.0014756103],"domain_scores_gemma":[0.89001775,0.079026476,0.017133193,0.0041814456,0.003906858,0.0057342597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014242161,0.0005554625,0.00050364103,0.0046236287,0.007327983,0.009375183,0.0011971422,0.0027300022,0.006909326],"category_scores_gemma":[0.06786859,0.0005166287,0.00040476597,0.003034345,0.011192898,0.01346965,0.010270813,0.0018035381,0.0005755024],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027951816,0.00045667135,0.09435706,0.0008587129,0.00013725026,0.0017380175,0.5084818,0.00159829,0.00096045446,0.16377598,0.013109033,0.21424715],"study_design_scores_gemma":[0.00025188795,0.0009590285,0.07649319,0.0017888738,0.00019050034,0.0039236615,0.43801385,0.0049918853,0.0009992447,0.30592054,0.16619508,0.00027237032],"about_ca_topic_score_codex":0.0020086926,"about_ca_topic_score_gemma":0.0014021244,"teacher_disagreement_score":0.014242161,"about_ca_system_score_codex":0.002666347,"about_ca_system_score_gemma":0.0032843626,"threshold_uncertainty_score":0.0753206},"labels":[],"label_agreement":null},{"id":"W2160715838","doi":"10.1145/1806799.1806828","title":"Using information fragments to answer the questions developers ask","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":158,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Ask price; Computer science; World Wide Web; Information retrieval; Business","score_opus":0.02240552777286546,"score_gpt":0.29504162165140296,"score_spread":0.2726360938785375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160715838","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16935249,0.0010963188,0.7830829,0.004378689,0.000148983,0.0035881011,0.0024855093,0.010153864,0.025713172],"genre_scores_gemma":[0.3462335,0.0007810066,0.6375002,0.0008951804,0.000085517095,0.0025727407,0.0045020934,0.0008155722,0.0066142078],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9916317,0.004785803,0.00067439396,0.00090398017,0.0017016829,0.00030247655],"domain_scores_gemma":[0.88671345,0.094941705,0.004857002,0.006352411,0.005884356,0.0012511138],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010220972,0.0015994399,0.00061555725,0.0043897806,0.001387508,0.0022079549,0.0018747521,0.0025187356,0.008053333],"category_scores_gemma":[0.08856854,0.0008823802,0.00066281774,0.0025439784,0.001598614,0.008717963,0.0032206068,0.0012219207,0.003032583],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019389561,0.0006870509,0.03572516,0.0026791387,0.0001008112,0.001248304,0.13906288,0.005491517,0.029422347,0.02621815,0.04610876,0.71131694],"study_design_scores_gemma":[0.001133438,0.0033193436,0.060664963,0.0039026362,0.0004343787,0.0051972317,0.079210296,0.11473771,0.048280656,0.13141952,0.5507795,0.00092028326],"about_ca_topic_score_codex":0.0065109935,"about_ca_topic_score_gemma":0.005991358,"teacher_disagreement_score":0.010220972,"about_ca_system_score_codex":0.0017795169,"about_ca_system_score_gemma":0.001618878,"threshold_uncertainty_score":0.05405432},"labels":[],"label_agreement":null},{"id":"W2160749579","doi":"10.1109/vissoft.2014.32","title":"Information Visualization for Agile Software Development","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Agile software development; Visualization; Computer science; Software development; Agile Unified Process; Information visualization; Software visualization; Software; Knowledge sharing; Agile usability engineering; Knowledge management; Software engineering; Software development process; Data science; Software construction; Data mining","score_opus":0.013037319090069697,"score_gpt":0.2653959247751804,"score_spread":0.2523586056851107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160749579","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018572424,0.044753093,0.83172774,0.022708192,0.0012857164,0.0005261314,0.0012815592,0.007217793,0.07192727],"genre_scores_gemma":[0.1762297,0.018256616,0.7967236,0.00079456053,0.00036809847,0.000717441,0.0010025955,0.0007981711,0.005109272],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.992724,0.0050135464,0.00048368692,0.00037270144,0.0012530136,0.00015312947],"domain_scores_gemma":[0.9743217,0.018899385,0.0012696313,0.0026314135,0.0023150055,0.0005628654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077843363,0.00081918086,0.0005648785,0.004334061,0.0017165641,0.0065166247,0.0011137192,0.0016227848,0.0120876925],"category_scores_gemma":[0.028282017,0.00068329234,0.0008636173,0.006719001,0.0018728488,0.005865936,0.0031555865,0.002329618,0.0015822464],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016355407,0.00011716516,0.00263983,0.002673611,0.000090716945,0.00044347317,0.006536463,0.005958658,0.004460059,0.18638669,0.044468228,0.7460616],"study_design_scores_gemma":[0.00010072462,0.00022079077,0.004967174,0.0052369065,0.000114339346,0.0013831401,0.0042571854,0.026073502,0.005401013,0.37505648,0.5769777,0.0002110118],"about_ca_topic_score_codex":0.0026376585,"about_ca_topic_score_gemma":0.0022178644,"teacher_disagreement_score":0.0120876925,"about_ca_system_score_codex":0.0018579625,"about_ca_system_score_gemma":0.0031223674,"threshold_uncertainty_score":0.041167974},"labels":[],"label_agreement":null},{"id":"W2160767434","doi":"10.1007/978-3-319-11164-3_1","title":"First International Competition on Software for Runtime Verification","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; University of Waterloo","funders":"","keywords":"Computer science; Java; Competition (biology); Runtime verification; Event (particle physics); Software; Software engineering; Process (computing); Operating system; Programming language; Formal verification","score_opus":0.016734453758752005,"score_gpt":0.2545581799950946,"score_spread":0.2378237262363426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160767434","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0069479854,0.034077525,0.13903514,0.015130856,0.048114058,0.00032941138,0.0021677758,0.005925378,0.74827194],"genre_scores_gemma":[0.022314116,0.011787857,0.028369792,0.0019209813,0.0041169566,0.00015196379,0.004119162,0.004855036,0.92236406],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99740773,0.00040444135,0.000108546184,0.00027651942,0.0014487284,0.00035406303],"domain_scores_gemma":[0.9975969,0.0004965174,0.000058092875,0.00036500074,0.0009879147,0.00049562374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036732575,0.0017105932,0.0012190406,0.0028229991,0.0012277873,0.004486636,0.0022999516,0.0020867346,0.09634419],"category_scores_gemma":[0.003645713,0.00069231895,0.0014216217,0.0029196965,0.0010087515,0.0041390085,0.0036834036,0.0033786544,0.042307295],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011061931,0.0000856346,0.00011411006,0.0002896495,0.00002503982,0.00007261368,0.000115851115,0.0014104628,0.002809684,0.062292926,0.6243999,0.3082735],"study_design_scores_gemma":[0.000017356093,0.00006446468,0.0002509168,0.00017635037,0.000013783182,0.0001338777,0.000035687935,0.0016289628,0.002138358,0.018411797,0.97711086,0.000017657208],"about_ca_topic_score_codex":0.0013847373,"about_ca_topic_score_gemma":0.0029577876,"teacher_disagreement_score":0.09634419,"about_ca_system_score_codex":0.0020751206,"about_ca_system_score_gemma":0.003758421,"threshold_uncertainty_score":0.3223034},"labels":[],"label_agreement":null},{"id":"W2160841184","doi":"10.5555/2664446.2664451","title":"A linked data platform for mining software repositories","year":2012,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software versioning; Software; World Wide Web; Cloud computing; Software mining; BitTorrent tracker; Database; Software engineering; Software development; Data science; Software construction; Operating system","score_opus":0.05717894938690845,"score_gpt":0.3065874026592379,"score_spread":0.24940845327232944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160841184","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01223419,0.00092918915,0.5841338,0.0013581768,0.0002525233,0.0033408836,0.27210966,0.11557677,0.010064783],"genre_scores_gemma":[0.03464387,0.0007687688,0.47233188,0.00039523304,0.00009766586,0.0032980796,0.48077345,0.0038027407,0.0038882955],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99225956,0.001260577,0.0016189789,0.0015612367,0.0029942144,0.00030540256],"domain_scores_gemma":[0.98044854,0.0051992326,0.0024416954,0.0074548046,0.0032092175,0.0012464182],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076530334,0.0016259234,0.0013581993,0.019198036,0.0025359676,0.005473047,0.003818718,0.0021559289,0.0073091793],"category_scores_gemma":[0.031804666,0.0015944192,0.0023374956,0.021194346,0.0007974037,0.00890964,0.008550226,0.0029873874,0.006393253],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014404076,0.0010505414,0.029736903,0.0048840973,0.0014132732,0.0016804414,0.002915089,0.032135494,0.0148472795,0.11898215,0.30751058,0.48340377],"study_design_scores_gemma":[0.0005185564,0.0003388406,0.018770913,0.0009799444,0.00036259153,0.00088438694,0.0009625295,0.13869174,0.019515572,0.14143911,0.67706376,0.00047210458],"about_ca_topic_score_codex":0.009962943,"about_ca_topic_score_gemma":0.010423034,"teacher_disagreement_score":0.019198036,"about_ca_system_score_codex":0.002054416,"about_ca_system_score_gemma":0.0049691745,"threshold_uncertainty_score":0.04047358},"labels":[],"label_agreement":null},{"id":"W2161272535","doi":"10.1145/633292.633495","title":"\"Bloat\"","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Word (group theory); Software; Term (time); World Wide Web; Programming language; Linguistics","score_opus":0.011844982852723502,"score_gpt":0.2466909815784295,"score_spread":0.234845998725706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161272535","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13054541,0.010183107,0.090866424,0.019241562,0.0042390237,0.0004069135,0.0020989834,0.0035001596,0.73891836],"genre_scores_gemma":[0.710333,0.0044633816,0.02713497,0.022866799,0.0011723527,0.00044260887,0.002243088,0.0019036008,0.2294402],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949239,0.0020314502,0.0002805831,0.000633016,0.0015025903,0.00062845444],"domain_scores_gemma":[0.993526,0.0014110852,0.0015880063,0.0014188269,0.0015969302,0.0004592435],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0026990601,0.00087194203,0.000417939,0.0023275842,0.004468251,0.0056266943,0.0012616082,0.002169316,0.021112284],"category_scores_gemma":[0.010444642,0.00034638616,0.0005873529,0.002855851,0.0072445273,0.008584642,0.0063718734,0.0024698935,0.00913881],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022396776,0.00007035317,0.008940988,0.00080237613,0.00006254638,0.0006393749,0.052221198,0.00023463361,0.0027209918,0.51345885,0.28169435,0.13893038],"study_design_scores_gemma":[0.000008898993,0.00005393409,0.0056825876,0.0002694211,0.000023596936,0.0012615315,0.009746653,0.00033643324,0.0007043256,0.019012041,0.9628502,0.00005054053],"about_ca_topic_score_codex":0.0068964604,"about_ca_topic_score_gemma":0.010203153,"teacher_disagreement_score":0.97888774,"about_ca_system_score_codex":0.001951056,"about_ca_system_score_gemma":0.001654792,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2161356945","doi":"10.1109/wpc.1994.341265","title":"Determining the usefulness of colour and fonts in a programming task","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Font; Task (project management); Computer science; Code (set theory); Null (SQL); Artificial intelligence; Programming language; Data mining; Engineering","score_opus":0.041716228530554365,"score_gpt":0.254363576011671,"score_spread":0.2126473474811166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161356945","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99062777,0.0003626528,0.0046887607,0.000103695565,0.000071041526,0.00024160302,0.00016978612,0.0003462147,0.0033883776],"genre_scores_gemma":[0.9666417,0.00038692902,0.02831353,0.00020151054,0.00008874141,0.00042296288,0.0003868971,0.0003108826,0.0032468324],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99701023,0.001323736,0.00043038765,0.0004957822,0.00059966376,0.000140283],"domain_scores_gemma":[0.85867876,0.12676096,0.0052326787,0.004072419,0.003799106,0.0014560685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053558126,0.00087774626,0.00068276835,0.00078244385,0.00041089242,0.0016057006,0.00049057207,0.0008886983,0.0054088645],"category_scores_gemma":[0.0538495,0.00056724646,0.00041612264,0.00060071587,0.0005898356,0.0020513292,0.0005544363,0.0008498548,0.00085918175],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.030217802,0.0060421177,0.07039313,0.00238972,0.0003682582,0.0006828557,0.007693022,0.003371963,0.65504694,0.0012771435,0.0044703768,0.2180466],"study_design_scores_gemma":[0.0029330705,0.05374889,0.65252435,0.00041567805,0.0014680016,0.0017344845,0.0030367686,0.01866513,0.25022727,0.0025113665,0.01221179,0.00052319333],"about_ca_topic_score_codex":0.0005187839,"about_ca_topic_score_gemma":0.00075691094,"teacher_disagreement_score":0.0054088645,"about_ca_system_score_codex":0.0002305388,"about_ca_system_score_gemma":0.00027818262,"threshold_uncertainty_score":0.028324604},"labels":[],"label_agreement":null},{"id":"W2161450348","doi":"10.1109/scam.2008.31","title":"From Indentation Shapes to Code Structures","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Indentation; Code (set theory); Block (permutation group theory); Computer science; Function (biology); Matching (statistics); Expression (computer science); Algorithm; Programming language; Geometry; Mathematics; Statistics","score_opus":0.030227774210830696,"score_gpt":0.2924741526337349,"score_spread":0.2622463784229042,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161450348","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94029367,0.00050603633,0.053059753,0.00022079721,0.00003083461,0.00008652779,0.0006763092,0.0010342188,0.0040917564],"genre_scores_gemma":[0.9833679,0.00008609194,0.015058402,0.00003148204,0.00001719402,0.00004708258,0.00048024277,0.00027978644,0.00063182967],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949351,0.001521244,0.00041994918,0.001221961,0.0016753835,0.00022632381],"domain_scores_gemma":[0.8586625,0.085128985,0.02687732,0.012839887,0.0146696195,0.001821699],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034074718,0.00058232446,0.00044028382,0.0045273765,0.0008634038,0.002451907,0.0007450073,0.0005865973,0.0017877717],"category_scores_gemma":[0.096243456,0.00063804747,0.00052063464,0.0038252538,0.0019758192,0.0038941144,0.0023974618,0.0013622248,0.00041059774],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068782194,0.00010824544,0.7656445,0.000451391,0.0001779613,0.0009876037,0.01363709,0.010093219,0.023176007,0.00906226,0.001637126,0.17433673],"study_design_scores_gemma":[0.000028257919,0.00031748565,0.87889934,0.00016170996,0.0001120614,0.002066829,0.0036514471,0.05803135,0.01694495,0.029263385,0.010290518,0.00023281659],"about_ca_topic_score_codex":0.0017364442,"about_ca_topic_score_gemma":0.002732535,"teacher_disagreement_score":0.0045273765,"about_ca_system_score_codex":0.00091939553,"about_ca_system_score_gemma":0.00056778523,"threshold_uncertainty_score":0.01802069},"labels":[],"label_agreement":null},{"id":"W2161538570","doi":"10.1109/ase.2006.72","title":"Using Decision Trees to Predict the Certification Result of a Build","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"IBM (Canada); University of Victoria","funders":"","keywords":"Certification; Computer science; Software engineering; Process (computing); Usability; IBM; Code (set theory); Programming language; Operating system","score_opus":0.04342492015173524,"score_gpt":0.30822756158988207,"score_spread":0.2648026414381468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161538570","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63933635,0.0016387241,0.34678537,0.0012944867,0.00021745764,0.00045137687,0.004254126,0.0031367762,0.0028853647],"genre_scores_gemma":[0.8857799,0.0003572226,0.103786595,0.0002159111,0.00009429436,0.00025097557,0.0080150245,0.000098560624,0.0014014604],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979506,0.0008314164,0.00016970302,0.00046411727,0.00031959702,0.0002647173],"domain_scores_gemma":[0.9760291,0.020169392,0.0011391134,0.00041298097,0.0015044624,0.00074504886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005480522,0.002061162,0.001211262,0.0044710897,0.00079562335,0.0015468113,0.0011362468,0.0016884232,0.0015439264],"category_scores_gemma":[0.01080652,0.00066139764,0.0014219114,0.0016654831,0.00040311457,0.0012908778,0.00068282813,0.0018186754,0.00075346127],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059268106,0.00061357114,0.058890123,0.00014275925,0.00021447061,0.00030044047,0.0001424686,0.8434646,0.001057001,0.0009110728,0.0042130896,0.08945766],"study_design_scores_gemma":[0.000018304887,0.00006311865,0.0018640613,0.000016320624,0.00002674297,0.000025266163,0.00002878722,0.99601483,0.00034432573,0.0012748952,0.0003136711,0.000009693906],"about_ca_topic_score_codex":0.015311098,"about_ca_topic_score_gemma":0.014921754,"teacher_disagreement_score":0.015311098,"about_ca_system_score_codex":0.0012950574,"about_ca_system_score_gemma":0.00161289,"threshold_uncertainty_score":0.030443966},"labels":[],"label_agreement":null},{"id":"W2161652852","doi":"10.1109/icsm.2004.1357829","title":"Developing a multi-objective decision approach to select source-code improving transformations","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Maintainability; Source code; Process (computing); Dependency (UML); Code refactoring; Transformation (genetics); Software engineering; Program slicing; Programming language; Software quality; Quality (philosophy); Heuristic; Software; Software development; Artificial intelligence","score_opus":0.028843878001149813,"score_gpt":0.2808676425912748,"score_spread":0.252023764590125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161652852","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011876517,0.00009415351,0.9862191,0.00018296436,0.000009987004,0.00017287317,0.000037611997,0.00010904648,0.0012977817],"genre_scores_gemma":[0.23242746,0.0001343739,0.76573884,0.00008839828,0.00002764481,0.00039794442,0.0001124977,0.000058474183,0.0010144559],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99617493,0.0018787584,0.00025116015,0.00057890674,0.0008215166,0.00029468208],"domain_scores_gemma":[0.99277955,0.0050945464,0.00059180195,0.00016854216,0.0011272158,0.00023820445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071530864,0.0016805937,0.0018270967,0.0034328864,0.00088966795,0.002430574,0.002403036,0.0017680867,0.0029891306],"category_scores_gemma":[0.00809182,0.001193534,0.0015960315,0.0020234876,0.0011399877,0.001966466,0.0015069348,0.0021295222,0.0003506884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007263328,0.000192856,0.0010821259,0.00020923042,0.00014391447,0.000103504986,0.00022391473,0.9159652,0.0020337082,0.015222478,0.00047478394,0.06427563],"study_design_scores_gemma":[0.000020608006,0.000057719906,0.00010335299,0.000012132267,0.000025756004,0.00001122579,0.000036559355,0.9934908,0.00079965114,0.005090212,0.00034097646,0.0000109439525],"about_ca_topic_score_codex":0.005014655,"about_ca_topic_score_gemma":0.006022517,"teacher_disagreement_score":0.0071530864,"about_ca_system_score_codex":0.0028217111,"about_ca_system_score_gemma":0.0035453879,"threshold_uncertainty_score":0.037829578},"labels":[],"label_agreement":null},{"id":"W2161881828","doi":"10.1016/s0965-9978(01)00050-3","title":"QEST nD: n-dimensional extension and generalisation of a software performance measurement model","year":2002,"lang":"en","type":"article","venue":"Advances in Engineering Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Dimension (graph theory); Viewpoints; Perspective (graphical); Computer science; Extension (predicate logic); Industrial engineering; Field (mathematics); Process (computing); Product (mathematics); Key (lock); Software; Engineering; Artificial intelligence; Mathematics","score_opus":0.0256490191154793,"score_gpt":0.23130710626097253,"score_spread":0.20565808714549325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161881828","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003642361,0.000039760485,0.99329174,0.00013417067,0.000052779433,0.000039602528,0.0002545075,0.00078378886,0.0017613475],"genre_scores_gemma":[0.43485966,0.0003340393,0.54969776,0.00031145388,0.00017660485,0.00057717884,0.0023077778,0.00067041226,0.011065041],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971852,0.00084715546,0.0002556872,0.0006008214,0.0009350482,0.00017603752],"domain_scores_gemma":[0.995747,0.0013462474,0.00034829436,0.0014459452,0.0009781483,0.00013439576],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032790052,0.00092592154,0.00096132327,0.0013362863,0.00058484136,0.0018217709,0.0023227537,0.0012963945,0.0044877753],"category_scores_gemma":[0.01407597,0.00060947164,0.0018560507,0.0014202995,0.0010940321,0.0038673019,0.0035996125,0.002638746,0.0015222294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036150363,0.0002600858,0.0055167372,0.00036556087,0.00016327159,0.00038215358,0.0005662647,0.4105666,0.0073096287,0.32691807,0.009714653,0.23787543],"study_design_scores_gemma":[0.000015639514,0.00005489648,0.0008394557,0.000022119957,0.000020594034,0.000083402265,0.000027683565,0.9083579,0.0013148665,0.08147366,0.0077598207,0.000029949719],"about_ca_topic_score_codex":0.0063983165,"about_ca_topic_score_gemma":0.0048686736,"teacher_disagreement_score":0.0063983165,"about_ca_system_score_codex":0.0013664477,"about_ca_system_score_gemma":0.0016297998,"threshold_uncertainty_score":0.017341197},"labels":[],"label_agreement":null},{"id":"W2162226791","doi":"10.1109/icpc.2013.6613853","title":"Improving the detection accuracy of evolutionary coupling","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Coupling (piping); Association rule learning; Precision and recall; Recall; Association (psychology); Data mining; Artificial intelligence; Machine learning; Engineering; Psychology; Cognitive psychology","score_opus":0.013560761261926145,"score_gpt":0.2416240410141819,"score_spread":0.22806327975225574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2162226791","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.74982566,0.001066714,0.24194132,0.0005902381,0.00008462816,0.00012208243,0.0004823842,0.002541231,0.003345679],"genre_scores_gemma":[0.91154575,0.00017401658,0.08684127,0.00008465343,0.000042012533,0.000036444664,0.00055107294,0.00006880694,0.0006559721],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99208987,0.0021036766,0.0009339777,0.0016193767,0.0027295924,0.00052356813],"domain_scores_gemma":[0.927504,0.046512112,0.0079176575,0.005971087,0.011074101,0.0010209916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00658298,0.0009942505,0.0012422008,0.008207729,0.00071556447,0.002759245,0.0017646705,0.0016312902,0.00068384706],"category_scores_gemma":[0.06771619,0.00055900076,0.00079405896,0.0036856825,0.0005834842,0.003609544,0.0026993144,0.0013166881,0.00063915853],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047476843,0.00037772776,0.48116067,0.00029274903,0.00032860166,0.0004158758,0.00081509684,0.016339872,0.018462492,0.0025346254,0.002418921,0.47637862],"study_design_scores_gemma":[0.000053156167,0.00039992182,0.11842687,0.00007430852,0.00030022196,0.0014675691,0.00053841,0.8317273,0.03693348,0.0064185695,0.0035667613,0.00009347401],"about_ca_topic_score_codex":0.0033056994,"about_ca_topic_score_gemma":0.0039198343,"teacher_disagreement_score":0.008207729,"about_ca_system_score_codex":0.00052282645,"about_ca_system_score_gemma":0.0010385306,"threshold_uncertainty_score":0.034814537},"labels":[],"label_agreement":null},{"id":"W2162385934","doi":"10.1109/csee.2003.1191381","title":"What cognitive activities are performed in student projects?","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Coding (social sciences); Software engineering; Implementation; Process (computing); Personal software process; Software development process; Software; Software Engineering Process Group; Software development; Software construction; Programming language","score_opus":0.03058972969866167,"score_gpt":0.30790868562248896,"score_spread":0.27731895592382727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2162385934","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99389404,0.0001591519,0.0028524664,0.0001166454,0.000010136701,0.000017898526,0.00005409255,0.000048151025,0.0028473255],"genre_scores_gemma":[0.99715185,0.00020501514,0.0018961895,0.00002184656,0.000008342911,0.000022824755,0.00010446741,0.000013110304,0.000576262],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99588865,0.0015007346,0.00042982478,0.0005204822,0.001238402,0.00042191296],"domain_scores_gemma":[0.97685605,0.009725901,0.006031341,0.0016037977,0.0032538164,0.0025291038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029452413,0.0003841822,0.00032714862,0.0022542838,0.00057723175,0.0028895673,0.00052120263,0.0007939227,0.00092313316],"category_scores_gemma":[0.032468475,0.00025797295,0.00038434667,0.0013362074,0.0006959533,0.0018687041,0.0012435385,0.000594195,0.0005382431],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046112612,0.00046914467,0.722117,0.00027396012,0.00013101836,0.0002715175,0.022028564,0.0010504109,0.007651338,0.0010042705,0.0007637239,0.24377787],"study_design_scores_gemma":[0.000027537375,0.0007369092,0.96254516,0.00015638644,0.00006867069,0.00094233383,0.0207517,0.004081236,0.0032387525,0.002704878,0.004651029,0.00009543733],"about_ca_topic_score_codex":0.001132754,"about_ca_topic_score_gemma":0.0016671956,"teacher_disagreement_score":0.0029452413,"about_ca_system_score_codex":0.000472062,"about_ca_system_score_gemma":0.00067787647,"threshold_uncertainty_score":0.015576124},"labels":[],"label_agreement":null},{"id":"W2162436321","doi":"10.1109/icpc.2011.26","title":"The NiCad Clone Detector","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":252,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Queen's University","funders":"","keywords":"Computer science; clone (Java method); Directory; XML; Plug-in; Scalability; Normalization (sociology); Operating system; Detector; Programming language; Extensibility","score_opus":0.03583715506802576,"score_gpt":0.24800179878230097,"score_spread":0.2121646437142752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2162436321","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041192144,0.0010410604,0.5326507,0.00041045452,0.00046718016,0.0006805922,0.01734763,0.38108736,0.02512283],"genre_scores_gemma":[0.15579669,0.0006760052,0.73768294,0.0007485015,0.00014897111,0.001998179,0.045241557,0.027146531,0.030560642],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9956905,0.00035180207,0.00023535767,0.001172024,0.002259424,0.00029087754],"domain_scores_gemma":[0.993624,0.001523802,0.0006386733,0.001725981,0.002188534,0.00029891427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027962283,0.0013102901,0.0015641467,0.0036784285,0.0011231771,0.0027340478,0.0025524003,0.00116491,0.009987331],"category_scores_gemma":[0.010238759,0.0010687206,0.0009992084,0.0020599538,0.0005997519,0.0019765962,0.002389035,0.0016543904,0.008920169],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001293548,0.00028379867,0.030204602,0.0010894876,0.0003033766,0.0010954809,0.00091802294,0.0071277777,0.09786918,0.017227462,0.39284146,0.44974577],"study_design_scores_gemma":[0.00021253411,0.00023846317,0.015269555,0.00021152612,0.00024813256,0.0027532578,0.00037654632,0.20356406,0.2954047,0.010166682,0.47115746,0.00039710614],"about_ca_topic_score_codex":0.0038512691,"about_ca_topic_score_gemma":0.004006899,"teacher_disagreement_score":0.009987331,"about_ca_system_score_codex":0.0014868914,"about_ca_system_score_gemma":0.0021112703,"threshold_uncertainty_score":0.033410907},"labels":[],"label_agreement":null},{"id":"W2162624365","doi":"10.1109/ase.2009.65","title":"Automatically Recommending Triage Decisions for Pragmatic Reuse Tasks","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; McGill University","funders":"","keywords":"Computer science; Reuse; Task (project management); Software engineering; Plan (archaeology); Process (computing); Triage; Software; Recommender system; Human–computer interaction; Code (set theory); Code reuse; Software system; World Wide Web; Programming language; Systems engineering","score_opus":0.03942043419170047,"score_gpt":0.33598196585163764,"score_spread":0.29656153165993715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2162624365","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29444122,0.0008603764,0.6442961,0.0018601065,0.00026744473,0.0018012975,0.0011261089,0.048062842,0.007284452],"genre_scores_gemma":[0.3796346,0.00037120664,0.6109105,0.00020625709,0.00010366971,0.0003913148,0.0019898016,0.0008928782,0.0054997164],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965281,0.0013174072,0.00034727418,0.0008337812,0.00079550094,0.00017808235],"domain_scores_gemma":[0.96820843,0.021734966,0.0022784534,0.0027329319,0.0042373426,0.00080784655],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005474057,0.0022624906,0.0014015976,0.004034041,0.0014295861,0.0023371796,0.0028852813,0.002921952,0.005032648],"category_scores_gemma":[0.041581526,0.0012208755,0.00081361266,0.0018283372,0.0005043247,0.0028704514,0.001005605,0.0018663113,0.004017856],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012097672,0.0012953694,0.034794237,0.0008314145,0.00022626044,0.00069664087,0.0037392802,0.021539075,0.026234196,0.002069713,0.025446104,0.88191795],"study_design_scores_gemma":[0.00045626814,0.0008464887,0.018944606,0.00029983596,0.0004400495,0.00096307596,0.0031177236,0.90540695,0.037143536,0.0061737313,0.025850639,0.00035715077],"about_ca_topic_score_codex":0.015446277,"about_ca_topic_score_gemma":0.03179238,"teacher_disagreement_score":0.015446277,"about_ca_system_score_codex":0.00082018063,"about_ca_system_score_gemma":0.0030006848,"threshold_uncertainty_score":0.030712724},"labels":[],"label_agreement":null},{"id":"W2163001363","doi":"10.1109/csmr.2001.914966","title":"Cohesion as changeability indicator in object-oriented systems","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Cohesion (chemistry); Computer science; Object-oriented programming; Software; Realm; Reliability engineering; Industrial engineering; Engineering; Programming language","score_opus":0.015881213165028978,"score_gpt":0.2641270239692474,"score_spread":0.24824581080421845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163001363","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8655178,0.0018765534,0.124194965,0.00034184777,0.00006698843,0.00016435215,0.00020748931,0.00070899207,0.006921025],"genre_scores_gemma":[0.9897883,0.00017161573,0.009475118,0.00001907418,0.000033838933,0.000048668713,0.0001589296,0.000036238704,0.00026815088],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9930936,0.0023997133,0.0005534521,0.00069112994,0.002966471,0.00029555077],"domain_scores_gemma":[0.93924874,0.040720623,0.011329879,0.0033184097,0.004450976,0.0009313403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045766896,0.0004439261,0.0006102203,0.0074002566,0.0006275244,0.0017510967,0.00046706395,0.0007607572,0.00061387254],"category_scores_gemma":[0.050222237,0.00035544,0.00040797703,0.0039360207,0.0018515183,0.0030749368,0.0012720096,0.00068166817,0.00013619779],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077636825,0.0005617922,0.46281394,0.0012311784,0.00041475505,0.001204878,0.012036872,0.0723493,0.0447567,0.058040638,0.0026342412,0.34317937],"study_design_scores_gemma":[0.000094884366,0.0020610266,0.66764665,0.00028771628,0.00032231607,0.0015731212,0.0028354821,0.2143658,0.024820512,0.07660459,0.009131057,0.0002568422],"about_ca_topic_score_codex":0.0014587346,"about_ca_topic_score_gemma":0.0007370917,"teacher_disagreement_score":0.0074002566,"about_ca_system_score_codex":0.0009811391,"about_ca_system_score_gemma":0.00039380544,"threshold_uncertainty_score":0.024204075},"labels":[],"label_agreement":null},{"id":"W2163094228","doi":"10.1109/icsm.2002.1167795","title":"Combining software quality predictive models: an evolutionary approach","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software quality; Quality (philosophy); Software evolution; Software; Software construction; Software development; Programming language","score_opus":0.062433129095407304,"score_gpt":0.2970967247553559,"score_spread":0.23466359565994863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163094228","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029752426,0.0005143629,0.9641111,0.0005362113,0.00003604138,0.00017179499,0.00004672471,0.00044835857,0.0043829847],"genre_scores_gemma":[0.40949616,0.0008716696,0.58507884,0.00031343524,0.000117204756,0.00043871367,0.00034451584,0.00014696174,0.003192393],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99633634,0.0015906094,0.0001696183,0.00056283735,0.0010628817,0.0002777333],"domain_scores_gemma":[0.9941918,0.0038405391,0.00033512028,0.0004106134,0.0010682044,0.00015366757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073866234,0.0020812736,0.0024089634,0.004018291,0.00079416676,0.0021689383,0.0038903665,0.0021888565,0.0017016261],"category_scores_gemma":[0.01606439,0.0014944196,0.0017471196,0.0029819019,0.0009618171,0.0027524554,0.002745736,0.0020452528,0.0005533171],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006222055,0.00023749526,0.004472117,0.00009053261,0.0003783406,0.00019632197,0.00021091901,0.80416864,0.0012171926,0.008705041,0.00068310066,0.17957811],"study_design_scores_gemma":[0.000010540705,0.000052686744,0.00039638855,0.00002541026,0.00006175268,0.000040443047,0.000035965368,0.99257576,0.0003627932,0.005778183,0.00064512226,0.000014987679],"about_ca_topic_score_codex":0.005664273,"about_ca_topic_score_gemma":0.004887467,"teacher_disagreement_score":0.0073866234,"about_ca_system_score_codex":0.0012729416,"about_ca_system_score_gemma":0.0011463243,"threshold_uncertainty_score":0.039064646},"labels":[],"label_agreement":null},{"id":"W2163099486","doi":"10.1109/ctgdsd.2013.6635240","title":"Tool usage within a globally distributed software development course and implications for teaching","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Class (philosophy); Software; Software engineering; Point (geometry); Software development; Course (navigation); Feature (linguistics); Data science; Artificial intelligence; Engineering; Programming language","score_opus":0.016037723549204453,"score_gpt":0.27920260428271687,"score_spread":0.2631648807335124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163099486","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99476707,0.00011691176,0.0023668415,0.00041200357,0.00001352379,0.000029524761,0.000014826697,0.000045018125,0.0022342578],"genre_scores_gemma":[0.99530727,0.00014546355,0.0027154875,0.00010094009,0.0000100213665,0.000043696244,0.00002462773,0.00003850991,0.0016140498],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9865283,0.0068448037,0.0006077722,0.0019463552,0.002549312,0.001523531],"domain_scores_gemma":[0.9570019,0.030073741,0.0024824643,0.001752388,0.004179838,0.004509727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008570786,0.0007255072,0.00073307246,0.002819895,0.003326643,0.006627955,0.0017215187,0.0015443725,0.0023513776],"category_scores_gemma":[0.036430318,0.00041644028,0.0004490727,0.0022823415,0.0027269556,0.0035945862,0.0045997985,0.002038858,0.0007290999],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005553272,0.004032138,0.16547666,0.00046374643,0.00005641281,0.0050409916,0.4240537,0.0017557252,0.01918716,0.0026814628,0.0022320412,0.37446463],"study_design_scores_gemma":[0.0001431304,0.005518229,0.2715091,0.0009759613,0.00014906375,0.0050228266,0.6387596,0.007762838,0.02571228,0.00833341,0.035739325,0.00037414258],"about_ca_topic_score_codex":0.001652572,"about_ca_topic_score_gemma":0.0035552816,"teacher_disagreement_score":0.008570786,"about_ca_system_score_codex":0.0027602438,"about_ca_system_score_gemma":0.002344835,"threshold_uncertainty_score":0.045327187},"labels":[],"label_agreement":null},{"id":"W2163309370","doi":"10.1109/icsm.2008.4658070","title":"An empirical study of the relationships between design pattern roles and class change proneness","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"Ministero dell'Università e della Ricerca","keywords":"Software design pattern; Computer science; Design pattern; Structural pattern; Eclipse; Empirical research; Software design; Software; Software development; Data science; Software engineering; Programming language; Mathematics","score_opus":0.2244818434539402,"score_gpt":0.3409468513609697,"score_spread":0.11646500790702949,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163309370","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99761224,0.000053991505,0.0013595468,0.0000643275,0.0000025295644,0.000034726512,0.00015322407,0.000014079986,0.00070546084],"genre_scores_gemma":[0.9985368,0.000023616501,0.0011030869,0.0000096253825,0.0000035143057,0.000032207852,0.00015549018,0.0000053234944,0.00013019113],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9917932,0.0038078187,0.00093486474,0.0013384047,0.0017272509,0.00039846404],"domain_scores_gemma":[0.38738,0.54569256,0.04392565,0.00875926,0.011405774,0.0028367247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009756874,0.0003019788,0.0002575154,0.0024355843,0.000464815,0.0011144775,0.00061359024,0.00068687875,0.002542002],"category_scores_gemma":[0.15897219,0.00033530057,0.00032057415,0.0022037989,0.0011277931,0.0022739018,0.00084171945,0.0015090118,0.000255301],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018558146,0.00030774437,0.9883456,0.00006347599,0.000058905578,0.000073787676,0.0013912997,0.00024644262,0.000579004,0.00021693528,0.000094772906,0.008436483],"study_design_scores_gemma":[0.000012923164,0.00040701785,0.9937157,0.000015444568,0.00002988938,0.0002696005,0.0016059919,0.0026052534,0.00070974586,0.0002877773,0.00032841452,0.000012292528],"about_ca_topic_score_codex":0.00095005264,"about_ca_topic_score_gemma":0.0017696831,"teacher_disagreement_score":0.009756874,"about_ca_system_score_codex":0.00057845854,"about_ca_system_score_gemma":0.00054111524,"threshold_uncertainty_score":0.05159992},"labels":[],"label_agreement":null},{"id":"W2163689760","doi":"10.1109/wcre.2004.37","title":"The small world of software reverse engineering","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Reverse engineering; Business process reengineering; Computer science; Software engineering; Data science; Social software engineering; Software maintenance; Software; Software development; Software construction; Engineering; Programming language; Manufacturing engineering","score_opus":0.014798473245051796,"score_gpt":0.22695862290868182,"score_spread":0.21216014966363003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163689760","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22117221,0.061858203,0.18286143,0.20903505,0.0035076037,0.00022891354,0.0018098358,0.0018684592,0.31765825],"genre_scores_gemma":[0.85955596,0.033707183,0.053498525,0.010981909,0.004751114,0.00034606218,0.0011249367,0.00085094967,0.035183422],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962287,0.0018854988,0.00009258281,0.00068096485,0.00094253937,0.00016969818],"domain_scores_gemma":[0.94218105,0.04514998,0.0031218722,0.005459809,0.002282235,0.0018050874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004923405,0.0004929871,0.00078433775,0.0044356766,0.003831033,0.010867566,0.0012938627,0.0026026517,0.013775463],"category_scores_gemma":[0.034056257,0.00054310693,0.00055917155,0.0051187123,0.008107768,0.022350987,0.0046327054,0.0035356057,0.003019215],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026107705,0.00011533003,0.0070114075,0.00080480915,0.00011579204,0.0015031468,0.0049363957,0.003228754,0.0018819153,0.6180773,0.072309196,0.28975487],"study_design_scores_gemma":[0.000031808737,0.000083568935,0.0037478209,0.00037222347,0.00003874622,0.0012370782,0.004232639,0.004008378,0.00083846063,0.62100685,0.3643505,0.00005180908],"about_ca_topic_score_codex":0.0012048694,"about_ca_topic_score_gemma":0.0017392162,"teacher_disagreement_score":0.013775463,"about_ca_system_score_codex":0.0011737384,"about_ca_system_score_gemma":0.0019441055,"threshold_uncertainty_score":0.04608351},"labels":[],"label_agreement":null},{"id":"W2163835670","doi":"10.1109/sess.1997.595954","title":"From software metrics to software measurement methods: a process model","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":87,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software metric; Software measurement; Software sizing; Verification and validation; Software quality; Software; Metric (unit); Process (computing); Data mining; Software development process; Goal-Driven Software Development Process; Software construction; Software development; Software engineering; Reliability engineering; Engineering; Programming language","score_opus":0.11320094386298554,"score_gpt":0.3423479141088234,"score_spread":0.22914697024583788,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163835670","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016051041,0.001246517,0.9791788,0.004547218,0.00014572406,0.00027268208,0.00011293754,0.00024242346,0.012648549],"genre_scores_gemma":[0.1449049,0.004816309,0.83759075,0.0017264445,0.00054098875,0.002193575,0.00041080045,0.0002788903,0.0075373813],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98557913,0.007410886,0.0009308214,0.001210909,0.004407252,0.00046087528],"domain_scores_gemma":[0.97436833,0.019755602,0.0011384572,0.0014306946,0.0027556357,0.00055130455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013255083,0.0016193297,0.0010953982,0.005303399,0.0014794794,0.008591897,0.0032858928,0.0065294495,0.0054663313],"category_scores_gemma":[0.035142437,0.0012323175,0.0015621678,0.0074013565,0.007239995,0.018047426,0.0042732838,0.0068283775,0.002700651],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000014185581,0.00005485041,0.00032393038,0.0002038037,0.000017840024,0.00010081609,0.000750773,0.008907743,0.00037044246,0.9644081,0.0018162341,0.023031194],"study_design_scores_gemma":[0.000030319967,0.00009857699,0.00032602635,0.0003486517,0.000027237647,0.00014723046,0.00020969496,0.04857812,0.0008467091,0.904283,0.04505324,0.000051258743],"about_ca_topic_score_codex":0.0044381283,"about_ca_topic_score_gemma":0.0024719974,"teacher_disagreement_score":0.013255083,"about_ca_system_score_codex":0.004957854,"about_ca_system_score_gemma":0.0064223493,"threshold_uncertainty_score":0.07010043},"labels":[],"label_agreement":null},{"id":"W2163837601","doi":"","title":"Revisiting the Impact of Classification Techniques on the Performance of Defect Prediction Models","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":264,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Replicate; Machine learning; Predictive modelling; Software; Artificial intelligence; Multivariate adaptive regression splines; Data mining; Multivariate statistics; Software bug; Set (abstract data type); Variety (cybernetics); Regression; Support vector machine; Software quality; Logistic regression; Regression analysis; Statistics; Software development; Mathematics","score_opus":0.077198300332788,"score_gpt":0.31477435560563966,"score_spread":0.23757605527285164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163837601","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9399507,0.008561662,0.03724997,0.0028968796,0.000896674,0.00020868237,0.0037465678,0.0026986105,0.0037903225],"genre_scores_gemma":[0.958965,0.0006948136,0.033160444,0.00033292818,0.00029537658,0.00007925631,0.005607225,0.00025833695,0.00060658425],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9787845,0.011497964,0.0015388515,0.0037065812,0.0035013994,0.00097063824],"domain_scores_gemma":[0.8083639,0.1448201,0.008287071,0.022615416,0.013905538,0.0020080856],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03136452,0.0024513411,0.0018664944,0.0038434896,0.001386161,0.003983231,0.0030772064,0.0030892498,0.0010804115],"category_scores_gemma":[0.11086161,0.00067014195,0.0025819587,0.003631491,0.0016624832,0.00595384,0.0025830716,0.0062038847,0.0015389089],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006148509,0.0025393486,0.31686428,0.0011507502,0.002186293,0.00034560816,0.000853087,0.24183522,0.009683012,0.0019326102,0.026305113,0.39015612],"study_design_scores_gemma":[0.00023488433,0.0013988495,0.04387443,0.00025565576,0.00045386396,0.00024344272,0.00064145867,0.93277335,0.01186155,0.0039122384,0.0042165434,0.0001337676],"about_ca_topic_score_codex":0.009672066,"about_ca_topic_score_gemma":0.0069589964,"teacher_disagreement_score":0.03136452,"about_ca_system_score_codex":0.0013581682,"about_ca_system_score_gemma":0.0017606657,"threshold_uncertainty_score":0.16587341},"labels":[],"label_agreement":null},{"id":"W2163952773","doi":"10.1109/cmpsac.2003.1245356","title":"Incremental transformation of procedural systems to object oriented platforms","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Legacy system; Programming language; Business process reengineering; Object-oriented programming; Program transformation; Source code; Process (computing); Software maintenance; Representation (politics); Object (grammar); Legacy code; Method; Model transformation; Software engineering; Software system; Software; Artificial intelligence; Engineering","score_opus":0.013538492387529938,"score_gpt":0.25350120308011137,"score_spread":0.23996271069258143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163952773","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045264214,0.00006424148,0.94369376,0.00013209786,0.000052841282,0.00019700832,0.000066796296,0.0056884573,0.004840573],"genre_scores_gemma":[0.3376428,0.00032093708,0.65194374,0.00015546047,0.000037901682,0.00020838453,0.0006781669,0.0017194381,0.0072931442],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993742,0.000105092484,0.00003826517,0.0000878734,0.00031987563,0.00007470432],"domain_scores_gemma":[0.99897695,0.0002608927,0.00008927526,0.00046295638,0.00016819875,0.000041744304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006129631,0.00036986385,0.00020527371,0.0006796113,0.00038420822,0.0010775534,0.0011023592,0.0004882507,0.0015815467],"category_scores_gemma":[0.0038853993,0.00034088077,0.00078357727,0.00040936703,0.00069176365,0.0009182522,0.0019673626,0.0009184495,0.0006077219],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001877975,0.00039427858,0.0039416715,0.00031485694,0.00011892468,0.0027941505,0.0027304513,0.12711573,0.19516858,0.21974209,0.0067321826,0.44075936],"study_design_scores_gemma":[0.00010457269,0.00036313204,0.0028583624,0.00011107761,0.00013516692,0.0015162865,0.00058315875,0.51918143,0.2057905,0.10775821,0.16150957,0.00008853033],"about_ca_topic_score_codex":0.001081743,"about_ca_topic_score_gemma":0.0010546594,"teacher_disagreement_score":0.0015815467,"about_ca_system_score_codex":0.00032697213,"about_ca_system_score_gemma":0.00071023736,"threshold_uncertainty_score":0.0052908063},"labels":[],"label_agreement":null},{"id":"W2164059708","doi":"10.1109/wse.2005.11","title":"REGoLive: Building a Web Site Comprehension Tool by Extending GoLive","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Program comprehension; Computer science; Comprehension; Class hierarchy; Software engineering; World Wide Web; Software; Compiler; Human–computer interaction; Programming language; Software system; Object-oriented programming","score_opus":0.011376693324134653,"score_gpt":0.2700484828927164,"score_spread":0.2586717895685817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164059708","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01533777,0.00016333317,0.81679803,0.0002347547,0.00006116833,0.00040714716,0.00027642644,0.16000898,0.006712366],"genre_scores_gemma":[0.08454672,0.0002750984,0.89237916,0.0004485027,0.000069258036,0.00040563144,0.0016919969,0.012080511,0.008103119],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998054,0.00045527486,0.0002045894,0.0003185708,0.0008335955,0.00013390163],"domain_scores_gemma":[0.9936371,0.0031334753,0.0004326968,0.0014508651,0.0010713006,0.000274593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003304842,0.001294143,0.0007474013,0.0022142532,0.00031706868,0.0019660962,0.0028765053,0.0012858583,0.004293725],"category_scores_gemma":[0.010724432,0.0009361436,0.0011724199,0.0005608676,0.00084201916,0.0037507655,0.0023108632,0.002137971,0.0037725163],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002943137,0.0006086359,0.006319931,0.0010719601,0.00011120909,0.001701632,0.0023301246,0.0065408256,0.06812427,0.009588812,0.027588136,0.8757202],"study_design_scores_gemma":[0.00049464434,0.0012927181,0.009028914,0.0006536845,0.00033425773,0.007881139,0.0007451711,0.22238173,0.18781304,0.01617441,0.55269873,0.0005015496],"about_ca_topic_score_codex":0.00087096513,"about_ca_topic_score_gemma":0.0012972171,"teacher_disagreement_score":0.004293725,"about_ca_system_score_codex":0.00039536142,"about_ca_system_score_gemma":0.0009427844,"threshold_uncertainty_score":0.01747787},"labels":[],"label_agreement":null},{"id":"W2164067145","doi":"10.1109/ccece.2007.323","title":"Software Costs and Schedule Estimations Based on the Work Coordination Laws","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Duration (music); Workload; Schedule; COCOMO; Computer science; Estimation; Software; Work (physics); Operations research; Law; Software development; Software engineering; Industrial engineering; Systems engineering; Engineering; Software construction; Political science","score_opus":0.018163962171963588,"score_gpt":0.27070469010886916,"score_spread":0.2525407279369056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164067145","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29105765,0.0001805847,0.69717556,0.00026362413,0.000021500857,0.00008060411,0.00017411784,0.00023337148,0.01081294],"genre_scores_gemma":[0.94216895,0.000108233944,0.056470316,0.000014168921,0.0000107827345,0.000076434175,0.00012613829,0.000032163985,0.0009927878],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99775356,0.0008592133,0.00011096191,0.00021859315,0.00090760575,0.00014999951],"domain_scores_gemma":[0.9922733,0.005037703,0.0011525227,0.00064383453,0.0007739473,0.000118694654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025348302,0.0005338518,0.00047387916,0.0022206171,0.0003125595,0.0011178353,0.0007458729,0.00046831334,0.0014903005],"category_scores_gemma":[0.021595517,0.0003755111,0.00047366443,0.0011768215,0.0008833602,0.0018706198,0.0006776232,0.00048349221,0.0001538479],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008082048,0.00005327469,0.0068954644,0.000044319557,0.000026324762,0.00005885173,0.00013783012,0.8557473,0.0010424561,0.10320441,0.0005544247,0.03215459],"study_design_scores_gemma":[0.0000054857765,0.00003554863,0.0018271046,0.000009431638,0.0000054690613,0.000016972237,0.000031264008,0.9813038,0.00066862477,0.01556165,0.0005255366,0.00000907255],"about_ca_topic_score_codex":0.0072428747,"about_ca_topic_score_gemma":0.005266285,"teacher_disagreement_score":0.0072428747,"about_ca_system_score_codex":0.0026155259,"about_ca_system_score_gemma":0.0012919017,"threshold_uncertainty_score":0.018977046},"labels":[],"label_agreement":null},{"id":"W2164095450","doi":"10.1145/1370175.1370247","title":"4th international workshop on predictor models in SE (PROMISE 2008)","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"USable; Computer science; Verifiable secret sharing; Software engineering; Software; Data science; World Wide Web; Set (abstract data type); Programming language","score_opus":0.048839451147974694,"score_gpt":0.2844698911114727,"score_spread":0.235630439963498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164095450","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00927382,0.0202002,0.90079236,0.025052384,0.0060838335,0.0002738583,0.0043998663,0.008768452,0.025155162],"genre_scores_gemma":[0.096048325,0.022320297,0.76669896,0.007433992,0.0058731716,0.0010532879,0.023498042,0.004706125,0.07236775],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.993948,0.003774976,0.0003147585,0.0007474562,0.0009466284,0.00026821622],"domain_scores_gemma":[0.9802461,0.011377957,0.00042824732,0.004765011,0.0023520384,0.00083056797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018031176,0.0017345285,0.002571159,0.001916076,0.0010857075,0.0057178307,0.004147445,0.0031540752,0.042972885],"category_scores_gemma":[0.0331169,0.001129593,0.0042202203,0.0024626744,0.0013110209,0.007963901,0.006085515,0.0072676027,0.016927058],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004863243,0.0006307788,0.0045915954,0.0008318913,0.00056903454,0.0003284897,0.0007066146,0.025349874,0.0009811204,0.09583131,0.2958493,0.57384366],"study_design_scores_gemma":[0.000173562,0.00042875705,0.0044202274,0.0013236246,0.00036598748,0.0005607111,0.000628315,0.19421996,0.0018403524,0.2252943,0.57059854,0.00014557574],"about_ca_topic_score_codex":0.006034938,"about_ca_topic_score_gemma":0.0093606105,"teacher_disagreement_score":0.042972885,"about_ca_system_score_codex":0.0015023176,"about_ca_system_score_gemma":0.0037700117,"threshold_uncertainty_score":0.14375865},"labels":[],"label_agreement":null},{"id":"W2164219553","doi":"10.1109/iceccs.2011.29","title":"A Novel Approach Based on Gestalt Psychology for Abstracting the Content of Large Execution Traces for Program Comprehension","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"TRACE (psycholinguistics); Computer science; Gestalt psychology; Program comprehension; Initialization; Context (archaeology); Software; Visualization; Theoretical computer science; Software system; Data mining; Programming language","score_opus":0.20679845585572185,"score_gpt":0.364153214215057,"score_spread":0.15735475835933513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164219553","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005219988,0.00007781008,0.9928591,0.00017477742,0.00002204123,0.00006416104,0.00003844854,0.0007875537,0.0007561204],"genre_scores_gemma":[0.10329094,0.00020862203,0.89443845,0.0000966744,0.000045988065,0.00018524796,0.00009993894,0.00028904062,0.0013452185],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99923897,0.00021738943,0.000042423486,0.00023031035,0.00021690756,0.000054051805],"domain_scores_gemma":[0.99779165,0.0012006314,0.00025973536,0.0003940114,0.00021848957,0.00013538834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009875689,0.0013370126,0.0008088123,0.003288233,0.00079284207,0.0022538248,0.0014592217,0.0013255166,0.0044723144],"category_scores_gemma":[0.006295594,0.00062892993,0.0016021255,0.0015722563,0.0028531773,0.0049634674,0.0019141356,0.0020342283,0.00075731607],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037237155,0.00032742112,0.0031637743,0.0006945169,0.00022693361,0.00046070496,0.0041664196,0.03240233,0.0877042,0.1980548,0.0041034943,0.66832304],"study_design_scores_gemma":[0.00009632235,0.00043014824,0.0047006574,0.00013068494,0.00013156192,0.000896862,0.000762812,0.6136073,0.030184586,0.32633853,0.022523152,0.00019736131],"about_ca_topic_score_codex":0.0017929452,"about_ca_topic_score_gemma":0.0021410894,"teacher_disagreement_score":0.0044723144,"about_ca_system_score_codex":0.000712774,"about_ca_system_score_gemma":0.0011250204,"threshold_uncertainty_score":0.014961362},"labels":[],"label_agreement":null},{"id":"W2164283454","doi":"10.1109/icpc.2009.5090051","title":"Methods for selecting and improving software clustering algorithms","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Computer science; Data mining; Software; Software system; Algorithm; Machine learning; Programming language","score_opus":0.026651403651026927,"score_gpt":0.3534773811205526,"score_spread":0.3268259774695257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164283454","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00080416317,0.00017479624,0.99793774,0.000058407477,0.000025441617,0.000074340205,0.000016637907,0.000427843,0.00048055607],"genre_scores_gemma":[0.011202374,0.00024691693,0.9874019,0.00004479024,0.000034127206,0.00020375264,0.00011081399,0.00018489714,0.0005705102],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9788334,0.006877409,0.0021418426,0.0028453933,0.008709315,0.00059260876],"domain_scores_gemma":[0.9716742,0.013058368,0.00210168,0.003925454,0.00888823,0.00035209276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01651905,0.0022872286,0.0018956321,0.006835052,0.0015649236,0.0034381824,0.0044278586,0.0019499887,0.0030398634],"category_scores_gemma":[0.042518485,0.0013558794,0.0024977338,0.005110516,0.0018167063,0.0047382507,0.003257234,0.0027800656,0.0019706534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015106543,0.00014277607,0.0024562408,0.0011065863,0.00018724956,0.00014450468,0.00074149383,0.0462558,0.013641267,0.10375562,0.0067610014,0.8246565],"study_design_scores_gemma":[0.00025951504,0.00033470348,0.0016563114,0.00055881776,0.0003704296,0.0010489852,0.00054328283,0.6712728,0.057036713,0.17096491,0.09568841,0.00026512207],"about_ca_topic_score_codex":0.0015187627,"about_ca_topic_score_gemma":0.0021972803,"teacher_disagreement_score":0.01651905,"about_ca_system_score_codex":0.0019215881,"about_ca_system_score_gemma":0.0027506,"threshold_uncertainty_score":0.08736217},"labels":[],"label_agreement":null},{"id":"W2164658662","doi":"","title":"Context-Sensitive Ranking of Dependencies for Software Navigation","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Ranking (information retrieval); Computer science; Dependency graph; Dependency (UML); Task (project management); Information retrieval; Graph; Data mining; Context (archaeology); Software; Theoretical computer science; Artificial intelligence; Engineering; Programming language","score_opus":0.020046839705642944,"score_gpt":0.27575610510151394,"score_spread":0.255709265395871,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164658662","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57882416,0.0037294754,0.40468818,0.0006002558,0.00020832707,0.00095956295,0.0010151207,0.005398826,0.0045760577],"genre_scores_gemma":[0.77989584,0.00041007472,0.21781804,0.00007456394,0.00007735351,0.00020310181,0.0006737113,0.00013772011,0.0007095243],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99440634,0.0029002635,0.00034323326,0.0007202764,0.0014017886,0.00022817214],"domain_scores_gemma":[0.9575023,0.03279158,0.0023583567,0.0031456307,0.003154407,0.0010477429],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053489916,0.0012306084,0.0012998922,0.0059780544,0.0013203236,0.001624844,0.0016997748,0.0018704676,0.0018094125],"category_scores_gemma":[0.043940824,0.00055384735,0.000735603,0.0030366874,0.00066416775,0.0035397978,0.0010069971,0.0013712412,0.0006442856],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002105145,0.0016452061,0.047148425,0.00087195815,0.00044738277,0.00025921868,0.00059601443,0.21866708,0.0145303225,0.013752769,0.008903983,0.6910726],"study_design_scores_gemma":[0.00014414219,0.00067994354,0.009163062,0.000041477575,0.00013674416,0.00030934703,0.00013544472,0.97245455,0.005269095,0.009847735,0.001702721,0.00011572214],"about_ca_topic_score_codex":0.0068650004,"about_ca_topic_score_gemma":0.015356354,"teacher_disagreement_score":0.0068650004,"about_ca_system_score_codex":0.0013845764,"about_ca_system_score_gemma":0.0024305834,"threshold_uncertainty_score":0.028288543},"labels":[],"label_agreement":null},{"id":"W2164694023","doi":"10.1145/2492248.2492261","title":"On the relationship between use cases and test suites size","year":2013,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Test Management Approach; Test (biology); Software quality; Test case; Regression testing; Software; Software engineering; Perspective (graphical); Software metric; Manual testing; Reliability engineering; Software system; Software development; Software construction; Programming language; Engineering; Machine learning; Artificial intelligence","score_opus":0.05497642424938919,"score_gpt":0.2632104469901711,"score_spread":0.2082340227407819,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164694023","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97622114,0.0021781668,0.015616699,0.0008895068,0.000033125732,0.00015041896,0.00062130636,0.00022292792,0.00406673],"genre_scores_gemma":[0.99466866,0.0003080953,0.0037792786,0.00006771131,0.000042598258,0.00011361464,0.00064885744,0.00006135401,0.00030984494],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.91885686,0.03831661,0.008328865,0.0074484274,0.024402685,0.0026465016],"domain_scores_gemma":[0.012387013,0.9664352,0.01407738,0.0026947237,0.0035909934,0.00081472535],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.046497524,0.0015411319,0.0010900514,0.009793767,0.00093346316,0.0044471445,0.0024931522,0.003163751,0.003346143],"category_scores_gemma":[0.6578684,0.0012015558,0.0015057528,0.007911823,0.0028614672,0.009695254,0.0024183218,0.004127842,0.00069876394],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005692484,0.0003795302,0.9685528,0.00016010515,0.00047229152,0.000400323,0.000635153,0.008353389,0.000714697,0.0009698833,0.00027293514,0.01851954],"study_design_scores_gemma":[0.000060869417,0.0012070119,0.9356025,0.00012317416,0.0004241433,0.0018317759,0.00094402867,0.05409838,0.0015143306,0.003477922,0.00060406193,0.00011191149],"about_ca_topic_score_codex":0.0026490188,"about_ca_topic_score_gemma":0.0029207172,"teacher_disagreement_score":0.046497524,"about_ca_system_score_codex":0.0015835975,"about_ca_system_score_gemma":0.0017667548,"threshold_uncertainty_score":0.24590534},"labels":[{"model":"gemma","categories":["metaresearch"],"domain":"methods","study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"},{"model":"gpt","categories":[],"domain":null,"study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"split"},{"id":"W2164878215","doi":"10.1109/tai.1989.65337","title":"Use of the W/AGE CASE tool in artificial intelligence","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Executable; Windsor; Computer science; Programming language; Artificial intelligence; Grammar; Rule-based machine translation; Natural language processing; Software engineering; Linguistics","score_opus":0.09100729089344374,"score_gpt":0.3001701160456993,"score_spread":0.2091628251522556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164878215","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005858167,0.00037762633,0.9542875,0.0012456732,0.00017609724,0.00009216594,0.00014195265,0.0047993553,0.033021413],"genre_scores_gemma":[0.089000754,0.00077146024,0.89614844,0.00035226336,0.00007647624,0.00017419684,0.00038280844,0.00063007633,0.012463573],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99792445,0.00077880576,0.00019052866,0.00023414225,0.000759344,0.00011261922],"domain_scores_gemma":[0.9967643,0.0022590475,0.00012111077,0.00050829514,0.00022421547,0.00012309122],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029016098,0.000670556,0.0003930516,0.0032207526,0.0010456757,0.004132686,0.001751482,0.0022365043,0.008194879],"category_scores_gemma":[0.0058411346,0.0006571108,0.00091773923,0.0019811352,0.0023499161,0.0057543046,0.0025233475,0.0018414134,0.0020429096],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009659865,0.00011579541,0.0009517376,0.0002783554,0.000039788927,0.0026125142,0.0008970169,0.010712999,0.004502789,0.7908229,0.01720089,0.17176867],"study_design_scores_gemma":[0.00011416044,0.000104118284,0.0006115859,0.0003512019,0.000050055976,0.0030191808,0.00034856022,0.123045765,0.0211178,0.43862268,0.4125129,0.00010198378],"about_ca_topic_score_codex":0.0016613204,"about_ca_topic_score_gemma":0.0024228238,"teacher_disagreement_score":0.008194879,"about_ca_system_score_codex":0.0008439441,"about_ca_system_score_gemma":0.00076211704,"threshold_uncertainty_score":0.02741468},"labels":[],"label_agreement":null},{"id":"W2164928043","doi":"10.1109/icse.2007.83","title":"Supporting the Investigation and Planning of Pragmatic Reuse Tasks","year":2007,"lang":"en","type":"article","venue":"Proceedings/Proceedings - International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Computer science; Task (project management); Software engineering; Code reuse; Software; Code (set theory); Systems engineering; Programming language; Engineering","score_opus":0.030283404631478442,"score_gpt":0.3020614974955754,"score_spread":0.271778092864097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164928043","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052266855,0.00017442499,0.9320926,0.0006228239,0.000030701052,0.00078633416,0.00014353119,0.009453444,0.0044292165],"genre_scores_gemma":[0.20664236,0.00019060732,0.78942907,0.00010562809,0.000016599613,0.0005380771,0.00048791632,0.000805055,0.0017847278],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9837518,0.008625797,0.0013119458,0.0016530979,0.0037065742,0.0009507616],"domain_scores_gemma":[0.88429695,0.08819158,0.0073716957,0.012969359,0.005414444,0.0017560283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012928289,0.0018871884,0.0012135283,0.0033886835,0.0018118556,0.004331438,0.0029808627,0.002523423,0.0042550983],"category_scores_gemma":[0.07484123,0.001988427,0.0013364747,0.0013508538,0.0025322167,0.005272441,0.005083166,0.0023329463,0.0017028367],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000979355,0.0016243744,0.023807611,0.0022471512,0.00017133764,0.0020808983,0.02053342,0.07320711,0.04659031,0.048539672,0.009915762,0.770303],"study_design_scores_gemma":[0.00054663623,0.0009807253,0.0103231445,0.0013484196,0.0002616479,0.0016979579,0.008754837,0.70560133,0.061872452,0.11329891,0.09472777,0.00058622466],"about_ca_topic_score_codex":0.0045390003,"about_ca_topic_score_gemma":0.006952139,"teacher_disagreement_score":0.012928289,"about_ca_system_score_codex":0.0013764926,"about_ca_system_score_gemma":0.008777976,"threshold_uncertainty_score":0.06837219},"labels":[],"label_agreement":null},{"id":"W2165129291","doi":"10.1109/csmr.2012.39","title":"Using fuzzy code search to link code fragments in discussions to source code","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Traceability; Computer science; Source code; Code review; KPI-driven code analysis; Internal documentation; Software engineering; Documentation; Static program analysis; Requirements traceability; Code (set theory); Software maintenance; Fuzzy logic; Software evolution; Software; Software development; Information retrieval; Programming language; Software construction; Artificial intelligence","score_opus":0.07597034745385421,"score_gpt":0.35819960755293806,"score_spread":0.28222926009908383,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165129291","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3626954,0.0007668427,0.6139182,0.00055287103,0.000056890447,0.0008379294,0.0007296527,0.004039613,0.016402608],"genre_scores_gemma":[0.6259892,0.0003126914,0.36955428,0.00007971749,0.00001696542,0.00034557347,0.0006219359,0.0002251245,0.0028546448],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99636024,0.0013207955,0.00027311762,0.0005621443,0.0013419342,0.0001417782],"domain_scores_gemma":[0.96165293,0.030118903,0.0031250978,0.0019377128,0.0027369726,0.00042836033],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043036663,0.0006702403,0.0005122201,0.016498223,0.0014692862,0.0025494944,0.0010325949,0.0009522044,0.004841367],"category_scores_gemma":[0.034590174,0.00040160678,0.0005242437,0.007673481,0.0012627641,0.0045009986,0.0020227155,0.00064784364,0.00085920346],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011622645,0.0004362979,0.040087055,0.0015190045,0.00017773235,0.0005679906,0.021408029,0.012677459,0.03952562,0.027726134,0.0026811422,0.85203123],"study_design_scores_gemma":[0.00036040982,0.0016256776,0.13734676,0.001344776,0.00054414035,0.002874212,0.024349498,0.5373721,0.12847526,0.12047521,0.044557188,0.0006747284],"about_ca_topic_score_codex":0.009472483,"about_ca_topic_score_gemma":0.010566056,"teacher_disagreement_score":0.016498223,"about_ca_system_score_codex":0.0016315395,"about_ca_system_score_gemma":0.0016691881,"threshold_uncertainty_score":0.022760212},"labels":[],"label_agreement":null},{"id":"W2165314999","doi":"10.1109/iri-05.2005.1506459","title":"Isolation of software defects: extracting knowledge with confidence","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software maintenance; Software; Isolation (microbiology); Software engineering; Software development; Software bug; Reliability engineering; Scheduling (production processes); Data mining; Engineering","score_opus":0.01846949362946427,"score_gpt":0.27629728471978987,"score_spread":0.2578277910903256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165314999","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27258846,0.0010535619,0.7196544,0.000579926,0.000033517015,0.00025773037,0.0011136002,0.0013831381,0.0033356766],"genre_scores_gemma":[0.8980878,0.00036829454,0.09962881,0.00006515824,0.000047199144,0.000104809995,0.0013553356,0.000046698755,0.00029581843],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948106,0.0011281568,0.00080633565,0.0011367166,0.0017721184,0.0003461811],"domain_scores_gemma":[0.9436184,0.044729855,0.003157628,0.0045831627,0.0035217793,0.0003891481],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005678501,0.0010706485,0.0018352895,0.008866518,0.00057588343,0.0035635468,0.0017864616,0.002151075,0.0010150413],"category_scores_gemma":[0.0707407,0.00061235845,0.001560355,0.005864296,0.0008407768,0.0056143,0.0025810774,0.0016808393,0.0005694963],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093429495,0.0007435484,0.110502064,0.0006448978,0.00053794275,0.00074812263,0.0010083931,0.11229557,0.005704927,0.005844862,0.0016370069,0.7593984],"study_design_scores_gemma":[0.000061132094,0.00034394034,0.02542598,0.0001475652,0.0003776491,0.00045544532,0.00035082496,0.9403964,0.007389943,0.023717778,0.0012440856,0.00008920745],"about_ca_topic_score_codex":0.004692783,"about_ca_topic_score_gemma":0.0020389268,"teacher_disagreement_score":0.008866518,"about_ca_system_score_codex":0.00070119876,"about_ca_system_score_gemma":0.0010805365,"threshold_uncertainty_score":0.030031145},"labels":[],"label_agreement":null},{"id":"W2165423824","doi":"10.1109/icsm.2007.4362650","title":"System-level Usage Dependency Analysis of Object-Oriented Systems","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Dependency graph; Dependency (UML); Class (philosophy); Software system; Context (archaeology); Graph; Software; Object-oriented programming; Software engineering; Programming language; Theoretical computer science; Data mining; Artificial intelligence","score_opus":0.023324480874765477,"score_gpt":0.27695466722371453,"score_spread":0.25363018634894907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165423824","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50147,0.00035292577,0.49141943,0.00021874771,0.000013532104,0.00018643012,0.0006107838,0.0020010218,0.0037271664],"genre_scores_gemma":[0.90724903,0.00014441204,0.09093534,0.00004370511,0.000009091036,0.000081171376,0.0007212241,0.00022772979,0.0005882306],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985688,0.0004899912,0.00009017024,0.00015300006,0.0005800339,0.00011807266],"domain_scores_gemma":[0.9945227,0.0030576386,0.0006678491,0.00097612507,0.0006911661,0.00008449762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010090927,0.00039351126,0.00037049837,0.0026569006,0.0005305068,0.0007646256,0.00048850605,0.00039313216,0.0006533851],"category_scores_gemma":[0.0057096467,0.0003666056,0.00048646773,0.0017501061,0.00068514165,0.0015487382,0.00059409905,0.00059190916,0.00011918531],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029963808,0.00040578935,0.18147904,0.0005291598,0.00031057457,0.0015118746,0.0031211812,0.32939437,0.05151581,0.08157035,0.0041145403,0.34574765],"study_design_scores_gemma":[0.000015502408,0.00008171989,0.051546387,0.000038591876,0.00011959294,0.00036565785,0.00025050587,0.87437695,0.020301994,0.04475951,0.008093292,0.000050278988],"about_ca_topic_score_codex":0.008384671,"about_ca_topic_score_gemma":0.010015858,"teacher_disagreement_score":0.008384671,"about_ca_system_score_codex":0.0008644443,"about_ca_system_score_gemma":0.00075877015,"threshold_uncertainty_score":0.016671717},"labels":[],"label_agreement":null},{"id":"W2165598415","doi":"10.1109/ase.1999.802337","title":"A metric based technique for design flaws detection and correction","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal; Université du Québec à Montréal","funders":"","keywords":"Computer science; Metric (unit); Transformation (genetics); Software quality; Task (project management); Process (computing); Quality (philosophy); Software metric; Software; Data mining; Reliability engineering; Software development; Systems engineering; Programming language; Engineering","score_opus":0.02527752021142424,"score_gpt":0.2624811347567859,"score_spread":0.23720361454536168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165598415","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030596375,0.00023692781,0.99230266,0.00012701961,0.000060158065,0.00017181115,0.00009666645,0.0033458623,0.00059924106],"genre_scores_gemma":[0.045855727,0.00016841119,0.9524132,0.00005845729,0.000026370235,0.00023454003,0.00028265745,0.0003081351,0.0006524462],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98543787,0.0024483816,0.0012770888,0.0015548026,0.008994487,0.00028741706],"domain_scores_gemma":[0.9735801,0.0064745336,0.0052956003,0.0059250216,0.008387115,0.00033769832],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062091025,0.0021616009,0.0020984015,0.007414084,0.0010221165,0.0018049275,0.0026949483,0.0019379568,0.0016208513],"category_scores_gemma":[0.03491046,0.0009337238,0.0014529162,0.004971707,0.001413503,0.0027219767,0.0020509223,0.0028005969,0.0009685285],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018369895,0.0002484319,0.004947941,0.0007762138,0.00021217996,0.00028947293,0.0003990161,0.022633253,0.074193865,0.017865757,0.006739118,0.8715111],"study_design_scores_gemma":[0.00012659335,0.0009609082,0.010958898,0.00018436972,0.0003013336,0.0032168424,0.00018096776,0.74467975,0.1736631,0.024030147,0.0413569,0.00034018193],"about_ca_topic_score_codex":0.0020260687,"about_ca_topic_score_gemma":0.0027306513,"teacher_disagreement_score":0.007414084,"about_ca_system_score_codex":0.0012585747,"about_ca_system_score_gemma":0.0022437645,"threshold_uncertainty_score":0.03283727},"labels":[],"label_agreement":null},{"id":"W2165688098","doi":"10.1145/941566.941569","title":"Static analysis to support the evolution of exception structure in object-oriented systems","year":2003,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":170,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Exception handling; Java; Control flow; Programming language; Software engineering; Control flow analysis; Object-oriented programming; Robustness (evolution); Static analysis; Static program analysis; Source code; Program code; Information flow; Software; Software development; Programming paradigm; Procedural programming; Inductive programming","score_opus":0.04314179430167696,"score_gpt":0.31228285955602275,"score_spread":0.2691410652543458,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165688098","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027318181,0.00015744899,0.95927864,0.00020088335,0.000037750768,0.000083309074,0.00010649866,0.010904479,0.0019129142],"genre_scores_gemma":[0.3413603,0.00031530915,0.6530353,0.00020625655,0.000099042,0.00022336794,0.0006915707,0.001872323,0.002196506],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997719,0.0007013709,0.0002057338,0.00021318998,0.0009819233,0.0001787707],"domain_scores_gemma":[0.9905697,0.005771393,0.0011663233,0.0011946203,0.0011364796,0.00016143029],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003328451,0.00076416455,0.00078890874,0.003572659,0.0010987195,0.001657991,0.0013370706,0.0009290273,0.0017504118],"category_scores_gemma":[0.015866864,0.00089979876,0.0009369204,0.0018570016,0.0014186592,0.0036033103,0.0015075252,0.0016157221,0.0004043793],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038336116,0.0003612521,0.02077246,0.0006440838,0.00013252538,0.001562361,0.004110036,0.23977277,0.03475048,0.26977688,0.009671343,0.4180625],"study_design_scores_gemma":[0.000050889965,0.00008051828,0.0021770776,0.00014102776,0.00008478268,0.00025314838,0.00016434895,0.86456543,0.020205315,0.09576637,0.016437301,0.000073773874],"about_ca_topic_score_codex":0.0043627447,"about_ca_topic_score_gemma":0.0046367706,"teacher_disagreement_score":0.0043627447,"about_ca_system_score_codex":0.0012610976,"about_ca_system_score_gemma":0.0021773654,"threshold_uncertainty_score":0.017602801},"labels":[],"label_agreement":null},{"id":"W2165739648","doi":"10.1109/icpc.2008.41","title":"NICAD: Accurate Detection of Near-Miss Intentional Clones Using Flexible Pretty-Printing and Code Normalization","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":519,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Parsing; Normalization (sociology); Agile software development; Programming language; Precision and recall; Transformation (genetics); Source code; Disk formatting; Code generation; Rule-based machine translation; Compiler; Code (set theory); Artificial intelligence; Natural language processing; Software engineering; Operating system","score_opus":0.0470208530369328,"score_gpt":0.29321803746072955,"score_spread":0.24619718442379676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165739648","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058868423,0.0004647537,0.9097724,0.00015272331,0.00007281176,0.00011727997,0.00026103706,0.028153049,0.002137561],"genre_scores_gemma":[0.17431314,0.00018505735,0.81894654,0.00015833127,0.00002461302,0.00015804765,0.00095522765,0.0015135505,0.0037453966],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966757,0.0005355753,0.00024444543,0.0007269615,0.001684431,0.0001328905],"domain_scores_gemma":[0.9915549,0.0032503938,0.0009122303,0.0022065325,0.0019207743,0.00015506303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002086135,0.00088354794,0.0012164158,0.002579856,0.0007649033,0.0018124742,0.0021087865,0.0011488365,0.001463324],"category_scores_gemma":[0.013768572,0.00052288605,0.00066328095,0.0017049401,0.00096386735,0.0022541077,0.0016944743,0.0011543878,0.0011259362],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053107075,0.00015657335,0.019544821,0.00035605612,0.00009840062,0.0005238942,0.00075757666,0.013232932,0.119207986,0.00791916,0.0066826423,0.83098894],"study_design_scores_gemma":[0.00008508107,0.00029864258,0.013541996,0.00007692396,0.000171523,0.0033399751,0.0002826367,0.4941291,0.44689423,0.009192022,0.03175686,0.000230951],"about_ca_topic_score_codex":0.0022711968,"about_ca_topic_score_gemma":0.0034687517,"teacher_disagreement_score":0.002579856,"about_ca_system_score_codex":0.0010649423,"about_ca_system_score_gemma":0.0017328548,"threshold_uncertainty_score":0.011032641},"labels":[],"label_agreement":null},{"id":"W2165754344","doi":"10.1109/euromicro.2006.20","title":"Analyzing Change Impact in Object-Oriented Systems","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Université de Montréal","funders":"","keywords":"Change impact analysis; Computer science; Object (grammar); Object-oriented programming; Software; Software maintenance; Software development; Industrial engineering; Reliability engineering; Artificial intelligence; Engineering","score_opus":0.021701528759515004,"score_gpt":0.28472872479648437,"score_spread":0.26302719603696934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165754344","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9712579,0.00010428216,0.027312987,0.0000394288,0.000006374692,0.00008713732,0.000049822807,0.00021294162,0.00092911377],"genre_scores_gemma":[0.99189246,0.00004621979,0.007784178,0.0000075075404,0.0000047392746,0.000030915035,0.00008511245,0.000011367183,0.00013753932],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970355,0.0011498433,0.00016728534,0.00023165057,0.001248442,0.00016735539],"domain_scores_gemma":[0.97539914,0.019825801,0.0018475342,0.0011667882,0.0015082443,0.000252494],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026884126,0.00055542117,0.00037880777,0.0033573348,0.0004064915,0.0006221051,0.0005046069,0.000512906,0.00077032694],"category_scores_gemma":[0.019188402,0.00024003655,0.00041293862,0.0019173238,0.00063970184,0.0016698194,0.00062050245,0.00046696022,0.000102955826],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011701668,0.0018910378,0.48777935,0.0005971961,0.00035215254,0.00069749425,0.0024764496,0.21298428,0.038230393,0.0041891118,0.00040575478,0.24922657],"study_design_scores_gemma":[0.000042800897,0.0014003313,0.29384077,0.000023144534,0.00015140092,0.0002376999,0.00075606204,0.66725945,0.030644806,0.0049434537,0.0006533086,0.000046672296],"about_ca_topic_score_codex":0.002774168,"about_ca_topic_score_gemma":0.0018166561,"teacher_disagreement_score":0.0033573348,"about_ca_system_score_codex":0.0007838454,"about_ca_system_score_gemma":0.00032891775,"threshold_uncertainty_score":0.014217854},"labels":[],"label_agreement":null},{"id":"W2165879825","doi":"10.1145/1368088.1368218","title":"Clonetracker","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":91,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Code refactoring; clone (Java method); Computer science; Software maintenance; Eclipse; Software evolution; Software engineering; Source code; Code reuse; Programming language; Reuse; Code (set theory); Software; Software development; Track (disk drive); Operating system; Software construction; Engineering","score_opus":0.02766108046267349,"score_gpt":0.25347726381612096,"score_spread":0.22581618335344747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165879825","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009786294,0.0013164916,0.38341388,0.00067730324,0.00097582326,0.00067173794,0.007482232,0.55950075,0.036175445],"genre_scores_gemma":[0.098592855,0.0023315717,0.55972034,0.0025471267,0.00045371143,0.0018510446,0.05184946,0.13379341,0.14886045],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966074,0.0003218183,0.0003488173,0.0009513552,0.0014955867,0.00027492],"domain_scores_gemma":[0.9912874,0.0024428144,0.0007060463,0.003054948,0.002144941,0.00036381496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003119903,0.0014773837,0.0012471474,0.0027220852,0.0013428873,0.0032680735,0.0038389752,0.002855012,0.03410473],"category_scores_gemma":[0.017367851,0.0016851752,0.0017510827,0.0017439256,0.00091006607,0.005541909,0.0042068814,0.0030756798,0.025111608],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011906747,0.00027059444,0.0075412956,0.0013899829,0.00020031718,0.0009369978,0.0014028574,0.0022010917,0.03154322,0.02517432,0.34515235,0.58299637],"study_design_scores_gemma":[0.00024966543,0.00018229873,0.0029502837,0.00023158442,0.00015508378,0.0014835652,0.00013506616,0.015090615,0.048008304,0.011451696,0.91987056,0.00019116425],"about_ca_topic_score_codex":0.0036662905,"about_ca_topic_score_gemma":0.004572638,"teacher_disagreement_score":0.03410473,"about_ca_system_score_codex":0.001002618,"about_ca_system_score_gemma":0.002271227,"threshold_uncertainty_score":0.114091694},"labels":[],"label_agreement":null},{"id":"W2165911035","doi":"10.1109/wcre.2005.27","title":"RETR: Reverse Engineering to Requirements","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); Queen's University; University of Toronto","funders":"","keywords":"Reverse engineering; Computer science; Software engineering; Software requirements; Software; Systems engineering; Software system; Software construction; Engineering; Programming language","score_opus":0.015843368315372055,"score_gpt":0.2519214410454726,"score_spread":0.23607807273010054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165911035","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016203242,0.0002944468,0.97180355,0.00071996247,0.00016280045,0.00028422242,0.0006886345,0.015953733,0.008472393],"genre_scores_gemma":[0.025480662,0.0006624847,0.9485272,0.00069045083,0.0001462671,0.00045468094,0.004054655,0.004733275,0.01525031],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98697287,0.0052683693,0.00084032246,0.0014893196,0.004971913,0.0004571918],"domain_scores_gemma":[0.97666556,0.010789514,0.001293708,0.007376058,0.0035381215,0.00033694407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067909774,0.002046058,0.0009882518,0.0038088593,0.0010401922,0.0035149471,0.002865263,0.0019286998,0.023474451],"category_scores_gemma":[0.03646194,0.0013396575,0.0025953366,0.0026391703,0.002245264,0.0057332846,0.0049748947,0.0048255916,0.016875748],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017873489,0.00020180957,0.000934796,0.0013242122,0.00010446414,0.0007972081,0.0013001492,0.01032128,0.009390035,0.20675887,0.083640106,0.6850483],"study_design_scores_gemma":[0.0001649058,0.0002816016,0.0007675379,0.00063558563,0.000102971746,0.0026319183,0.0007028445,0.102275014,0.03604743,0.2849946,0.5712738,0.000121690886],"about_ca_topic_score_codex":0.0017176658,"about_ca_topic_score_gemma":0.002262471,"teacher_disagreement_score":0.023474451,"about_ca_system_score_codex":0.0011745397,"about_ca_system_score_gemma":0.0021551596,"threshold_uncertainty_score":0.078529894},"labels":[],"label_agreement":null},{"id":"W2165929464","doi":"10.1109/icsm.2009.5306310","title":"What's hot and what's not: Windowed developer topic analysis","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":104,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Topic model; Commit; Focus (optics); Set (abstract data type); Data science; Search engine indexing; Timeline; Undo; World Wide Web; Information retrieval; Programming language; Database","score_opus":0.016487448542945173,"score_gpt":0.26648053634480534,"score_spread":0.24999308780186016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165929464","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4972103,0.0033231,0.48349997,0.0013693228,0.0002150103,0.00031231958,0.005068612,0.0040050824,0.0049962844],"genre_scores_gemma":[0.82818574,0.0008407524,0.16627014,0.00008802094,0.00012222501,0.00031587807,0.002227186,0.00036127513,0.0015888596],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990357,0.00032430427,0.0000775139,0.00026857987,0.00019836615,0.00009543442],"domain_scores_gemma":[0.99173915,0.0052338555,0.0010578629,0.00065985136,0.0009882395,0.00032099447],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00295205,0.00040685185,0.00044983998,0.0051465207,0.00059940876,0.0017456964,0.0005713108,0.00052536867,0.0011977279],"category_scores_gemma":[0.011870049,0.0002866135,0.0005098085,0.004419975,0.00039728353,0.0026696445,0.00091487385,0.00073848595,0.00041636865],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013874919,0.00024270151,0.1315624,0.0006323319,0.0003550834,0.00071233633,0.02451286,0.020637324,0.023748351,0.020282257,0.024463795,0.751463],"study_design_scores_gemma":[0.0001480164,0.0003632347,0.20463054,0.00034409462,0.00041889233,0.0013172908,0.012027552,0.6145795,0.01815146,0.066480234,0.0812129,0.00032633636],"about_ca_topic_score_codex":0.008682476,"about_ca_topic_score_gemma":0.007818936,"teacher_disagreement_score":0.008682476,"about_ca_system_score_codex":0.00060599745,"about_ca_system_score_gemma":0.0006900063,"threshold_uncertainty_score":0.01726389},"labels":[],"label_agreement":null},{"id":"W2165995531","doi":"10.1109/wcre.2012.20","title":"TRIS: A Fast and Accurate Identifiers Splitting and Expansion Algorithm","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Identifier; Computer science; Program comprehension; Tris; Unique identifier; Theoretical computer science; Representation (politics); Algorithm; Programming language; Software","score_opus":0.018869526102981385,"score_gpt":0.27547147508848463,"score_spread":0.25660194898550326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165995531","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017418424,0.00038070246,0.9552928,0.00015230704,0.00009530414,0.00018092032,0.0012612515,0.023110352,0.0021078333],"genre_scores_gemma":[0.048336394,0.00015393198,0.9415465,0.000116382595,0.000024630195,0.00017902271,0.00430147,0.0017406753,0.0036009753],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988763,0.00015792076,0.00012147124,0.00027697042,0.00046265856,0.00010476241],"domain_scores_gemma":[0.9983897,0.000527121,0.00016958092,0.00029065338,0.0005623131,0.000060603914],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009065498,0.001368705,0.0007660719,0.003066437,0.00070153835,0.00091260846,0.0018582792,0.0008460501,0.0066532823],"category_scores_gemma":[0.004408259,0.0006552525,0.0011380294,0.0025971874,0.00061907584,0.002326527,0.0020409662,0.0011823665,0.0039529433],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054485793,0.00011507334,0.0035796796,0.0004017021,0.00011676971,0.00028267963,0.0005286437,0.03429838,0.036398627,0.013100576,0.04716417,0.8634689],"study_design_scores_gemma":[0.00025883276,0.00039751362,0.0021817575,0.00010893116,0.00016940593,0.0006965649,0.00069144875,0.80469555,0.06924195,0.03425093,0.087191746,0.000115390365],"about_ca_topic_score_codex":0.004325218,"about_ca_topic_score_gemma":0.009030255,"teacher_disagreement_score":0.0066532823,"about_ca_system_score_codex":0.00080952287,"about_ca_system_score_gemma":0.0025092796,"threshold_uncertainty_score":0.022257447},"labels":[],"label_agreement":null},{"id":"W2166067341","doi":"10.1145/2145204.2145408","title":"Information needs for integration decisions in the release process of large-scale parallel development","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Process (computing); Process management; Work (physics); Computer science; Knowledge management; Context (archaeology); Key (lock); Information needs; Software; Scale (ratio); Information system; Software development process; Software development; Business; Engineering; World Wide Web; Computer security","score_opus":0.023031134683930716,"score_gpt":0.3002592292994753,"score_spread":0.2772280946155446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166067341","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9690657,0.00034925225,0.012526589,0.0050598723,0.000020491916,0.00022948139,0.00007660246,0.00010750726,0.012564487],"genre_scores_gemma":[0.99062294,0.00017300526,0.0079926085,0.00016203034,0.000014990423,0.00007969752,0.00010632656,0.000028190343,0.0008201754],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.97884244,0.013985897,0.0014412756,0.0009648608,0.003400413,0.0013650962],"domain_scores_gemma":[0.7314329,0.2325086,0.012111799,0.0040444974,0.015439325,0.004462941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028839098,0.00046760787,0.0004750186,0.0022601385,0.00314571,0.005711079,0.0013841452,0.0027146006,0.0035374407],"category_scores_gemma":[0.14991458,0.0010451818,0.00044059407,0.001305529,0.001715529,0.010272684,0.0023474735,0.001896086,0.00044951466],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020265798,0.0015273404,0.11310958,0.0018806227,0.000113309354,0.0072675357,0.5980935,0.005257653,0.0120696025,0.020856414,0.0058009573,0.23199691],"study_design_scores_gemma":[0.0002576293,0.0017324799,0.16480474,0.0015250732,0.0003637062,0.0035748763,0.6975107,0.0438425,0.0116851935,0.029344652,0.044884846,0.00047360544],"about_ca_topic_score_codex":0.0033963446,"about_ca_topic_score_gemma":0.0030380532,"teacher_disagreement_score":0.028839098,"about_ca_system_score_codex":0.00330438,"about_ca_system_score_gemma":0.0038310494,"threshold_uncertainty_score":0.15251756},"labels":[],"label_agreement":null},{"id":"W2166079152","doi":"10.1109/ccece.1995.526607","title":"Adapting the SIMAP productivity model to software maintenance","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Productivity; Computer science; Software maintenance; Software; Reliability engineering; Software engineering; Software development; Engineering; Operating system","score_opus":0.04689371735675383,"score_gpt":0.25501576589546027,"score_spread":0.20812204853870644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166079152","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043744728,0.00064144185,0.9230129,0.0014639539,0.00014362068,0.00022402455,0.0011263733,0.0014368088,0.028206075],"genre_scores_gemma":[0.63555604,0.0011238713,0.350086,0.00042010186,0.00025619767,0.00060486706,0.0020582776,0.00040266826,0.009491967],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99484116,0.0011899632,0.00043390918,0.0007450716,0.0023740747,0.00041589508],"domain_scores_gemma":[0.9879973,0.004797904,0.0013034418,0.0018530346,0.0037037672,0.00034449363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004679148,0.0012718108,0.0007685598,0.008027479,0.0009200274,0.004512752,0.0032395471,0.0011693755,0.004336496],"category_scores_gemma":[0.020569827,0.0004905578,0.0019256794,0.0062238397,0.0019007282,0.00637463,0.002822437,0.0019679721,0.0019560459],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030472042,0.00024296978,0.024311965,0.00038176242,0.00018248377,0.0004476623,0.0013910352,0.19571431,0.0018713265,0.44004565,0.00979346,0.32531253],"study_design_scores_gemma":[0.00004103422,0.00022069651,0.008324918,0.00011186622,0.00006023874,0.00041165884,0.000564546,0.5435012,0.0017239557,0.4169747,0.028006839,0.000058355883],"about_ca_topic_score_codex":0.007122386,"about_ca_topic_score_gemma":0.0042920723,"teacher_disagreement_score":0.008027479,"about_ca_system_score_codex":0.0039611747,"about_ca_system_score_gemma":0.0020142116,"threshold_uncertainty_score":0.028740466},"labels":[],"label_agreement":null},{"id":"W2166113239","doi":"10.1109/wse.2003.1234013","title":"Lessons learned in Web site architectures for public utilities","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Web site; Web modeling; Extensibility; Domain (mathematical analysis); Web application; World Wide Web; Web engineering; Web development; Architecture; Web standards; Web design; Software engineering; The Internet; Web application security","score_opus":0.07338286402822865,"score_gpt":0.3224089411746668,"score_spread":0.24902607714643815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166113239","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19065662,0.044176072,0.2360705,0.4117882,0.0025022966,0.00023490046,0.000284951,0.0009825757,0.113303915],"genre_scores_gemma":[0.6598585,0.047992367,0.23136741,0.014786168,0.0017250122,0.000188826,0.00036510782,0.0007272074,0.042989332],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9967907,0.0014148246,0.00019043621,0.00033149598,0.0009812426,0.00029116977],"domain_scores_gemma":[0.9876038,0.0068339924,0.0003746862,0.0012427313,0.0028455043,0.0010992247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007380043,0.0007051129,0.00035877732,0.0010475921,0.0016120357,0.00506658,0.002791901,0.0039398265,0.0036721844],"category_scores_gemma":[0.020410946,0.0006856103,0.00045089738,0.0011169635,0.004204917,0.0155158,0.0017313897,0.0059486325,0.0011728542],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000094924355,0.00043135957,0.008384471,0.0012959999,0.00006182981,0.0011959117,0.009128412,0.025010765,0.0030898121,0.4036207,0.04701521,0.5006707],"study_design_scores_gemma":[0.000075812044,0.0004289793,0.006761783,0.001750086,0.00005056706,0.0015850627,0.013806196,0.02914808,0.004691497,0.51308554,0.42844075,0.00017557856],"about_ca_topic_score_codex":0.016729958,"about_ca_topic_score_gemma":0.018460533,"teacher_disagreement_score":0.016729958,"about_ca_system_score_codex":0.0028227565,"about_ca_system_score_gemma":0.002072052,"threshold_uncertainty_score":0.039029837},"labels":[],"label_agreement":null},{"id":"W2166271660","doi":"10.1145/1882291.1882312","title":"Creating and evolving developer documentation","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":119,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Internal documentation; Computer science; World Wide Web; Technical documentation; Quality (philosophy); Software documentation; Code (set theory); Knowledge management; Software; Software development; Programming language; Set (abstract data type); Software development process","score_opus":0.010313125011201345,"score_gpt":0.2755106845546651,"score_spread":0.26519755954346375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166271660","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7176927,0.004180008,0.19130808,0.012090341,0.00082628086,0.0010132338,0.00033888145,0.0029830837,0.06956746],"genre_scores_gemma":[0.66328794,0.0031798915,0.30401543,0.0010326658,0.0002495285,0.00045816906,0.0008295657,0.0016125797,0.025334245],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9732015,0.012546952,0.002108772,0.0022029248,0.009050867,0.00088886486],"domain_scores_gemma":[0.8902292,0.045343354,0.012333173,0.021979393,0.025644416,0.0044703595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029698117,0.00054997305,0.00047243413,0.0044389837,0.004442808,0.007968246,0.002138773,0.0020650043,0.002405051],"category_scores_gemma":[0.12961115,0.00086394936,0.0005336138,0.0032481316,0.0026187182,0.010420118,0.005550851,0.0030021358,0.0011845862],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006658382,0.00026006295,0.06894133,0.0006665168,0.00005083433,0.0016108217,0.14627989,0.0014009189,0.011640938,0.020728067,0.017813936,0.73054016],"study_design_scores_gemma":[0.0000861695,0.0004847162,0.08971455,0.0029357707,0.00023007391,0.0065914798,0.07156223,0.0090186875,0.017053539,0.037263427,0.7647225,0.00033686505],"about_ca_topic_score_codex":0.0030615344,"about_ca_topic_score_gemma":0.0055790786,"teacher_disagreement_score":0.029698117,"about_ca_system_score_codex":0.002995448,"about_ca_system_score_gemma":0.007939643,"threshold_uncertainty_score":0.15706056},"labels":[],"label_agreement":null},{"id":"W2166414580","doi":"10.1109/msr.2007.7","title":"Determining Implementation Expertise from Bug Reports","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":110,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software engineering; Software bug; False positive paradox; Software; Software maintenance; Set (abstract data type); Product (mathematics); Code review; Software peer review; Data science; Software development; Static program analysis; Software construction; Programming language; Artificial intelligence","score_opus":0.022335965723631685,"score_gpt":0.33484647978740606,"score_spread":0.3125105140637744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166414580","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82709926,0.0021624477,0.14993376,0.00056187087,0.000112376845,0.0009824125,0.006885168,0.0022859785,0.009976715],"genre_scores_gemma":[0.87455434,0.0007579376,0.10946897,0.00013024855,0.00011146254,0.0005187369,0.012125331,0.00022195042,0.0021110761],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.96717304,0.008375038,0.004056154,0.0046932423,0.014496205,0.0012064311],"domain_scores_gemma":[0.6933117,0.18782753,0.03869571,0.024361132,0.053087637,0.0027164016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021452632,0.0011347277,0.0015745058,0.028855948,0.0010014551,0.0026213175,0.0014106077,0.0019421412,0.0012677093],"category_scores_gemma":[0.19486195,0.0007219418,0.0011368244,0.007938565,0.0007912993,0.004741967,0.0033952454,0.0016034066,0.0011434681],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004583575,0.00041039655,0.4650068,0.0016810346,0.0005507316,0.0006555779,0.0077628503,0.0077586058,0.008191017,0.0012637113,0.009344291,0.49691653],"study_design_scores_gemma":[0.00019522617,0.0009171458,0.8086644,0.0008723406,0.0009286721,0.0036002966,0.005607718,0.09787505,0.039301563,0.0076971343,0.033856396,0.00048410823],"about_ca_topic_score_codex":0.004879403,"about_ca_topic_score_gemma":0.006487314,"teacher_disagreement_score":0.028855948,"about_ca_system_score_codex":0.0009016074,"about_ca_system_score_gemma":0.001423359,"threshold_uncertainty_score":0.113453746},"labels":[],"label_agreement":null},{"id":"W2166700159","doi":"10.1145/1985793.1985844","title":"Identifying program, test, and environmental changes that affect behaviour","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Science Foundation","keywords":"Computer science; Test suite; Context (archaeology); XML; Affect (linguistics); Test (biology); Source code; Path (computing); Suite; Code (set theory); Test case; Programming language; Operating system; Psychology; Machine learning","score_opus":0.061317746701398,"score_gpt":0.2783450713997178,"score_spread":0.2170273246983198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166700159","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92536414,0.00024966535,0.06456396,0.00024254213,0.000036782043,0.00034881895,0.0014297921,0.001877461,0.0058869435],"genre_scores_gemma":[0.9720544,0.000122068814,0.024837086,0.000081502156,0.000012089239,0.00016828733,0.0012579843,0.00021753047,0.0012489931],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960646,0.0011213529,0.00025510776,0.0008135345,0.0014475578,0.0002978578],"domain_scores_gemma":[0.9786854,0.0114459675,0.0036114757,0.0022058599,0.003391548,0.00065972225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002193665,0.0007417801,0.00048416946,0.0018462103,0.00043465328,0.0013590216,0.0006386586,0.00094424473,0.0014166746],"category_scores_gemma":[0.025387846,0.00034504573,0.0004875403,0.0011786132,0.0007129463,0.0014278581,0.00079494173,0.00073983136,0.00059526775],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000849898,0.0006535407,0.8180015,0.0004768014,0.00021655446,0.0016684196,0.0019999756,0.016771771,0.035182968,0.0020276848,0.0017109538,0.12043987],"study_design_scores_gemma":[0.000062839645,0.00080975384,0.78947526,0.00012358864,0.0003423742,0.001648449,0.0021881952,0.15365477,0.035402276,0.00477956,0.0113756815,0.00013727415],"about_ca_topic_score_codex":0.0051343716,"about_ca_topic_score_gemma":0.010395788,"teacher_disagreement_score":0.0051343716,"about_ca_system_score_codex":0.0008557776,"about_ca_system_score_gemma":0.0012795761,"threshold_uncertainty_score":0.011601329},"labels":[],"label_agreement":null},{"id":"W2166779186","doi":"10.1109/suite.2009.5070012","title":"Search, stitch, view: Easing information integration in an IDE","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Interleaving; Variety (cybernetics); Software; Software engineering; Code (set theory); Information system; Information flow; World Wide Web; Programming language; Operating system; Engineering; Artificial intelligence","score_opus":0.020807709091370342,"score_gpt":0.29414694286656956,"score_spread":0.2733392337751992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166779186","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12833455,0.0014036315,0.7920282,0.0014080129,0.00021188027,0.0006135779,0.00046383936,0.06443467,0.011101596],"genre_scores_gemma":[0.20750202,0.000551632,0.77911204,0.0004633553,0.000080314465,0.00017221707,0.0010312239,0.0045772153,0.0065100407],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9944923,0.0018834898,0.0005173979,0.00060076645,0.0020775371,0.00042856295],"domain_scores_gemma":[0.9737587,0.015395003,0.0014492689,0.00626189,0.0020575393,0.0010776422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0094613675,0.00132483,0.001443259,0.0037430283,0.0016296966,0.0036402855,0.003027055,0.0024336074,0.0053842817],"category_scores_gemma":[0.039188325,0.0017336624,0.0013312827,0.0027349917,0.0019115017,0.016944809,0.007767541,0.0023338343,0.0028338432],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002001098,0.0011462203,0.015992006,0.0012131047,0.00021240993,0.0016366149,0.008408005,0.0064991764,0.03391258,0.01650425,0.023408376,0.8890662],"study_design_scores_gemma":[0.0030288144,0.0040150303,0.016261633,0.0012550477,0.0018621043,0.007909281,0.008967874,0.42917684,0.1951729,0.11823265,0.21321367,0.00090416666],"about_ca_topic_score_codex":0.0032734615,"about_ca_topic_score_gemma":0.0072638537,"teacher_disagreement_score":0.0094613675,"about_ca_system_score_codex":0.0007308118,"about_ca_system_score_gemma":0.0029108287,"threshold_uncertainty_score":0.050037146},"labels":[],"label_agreement":null},{"id":"W2167110832","doi":"10.1109/32.852742","title":"Validating the ISO/IEC 15504 measure of software requirements analysis process capability","year":2000,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":135,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Software measurement; Process (computing); Software engineering; Process capability; Software Engineering Process Group; Software; Software quality; Software development process; Reliability engineering; Systems engineering; Software development; Work in process; Engineering; Operations management; Operating system","score_opus":0.021389347909023807,"score_gpt":0.26770558990715365,"score_spread":0.24631624199812985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167110832","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85486466,0.0004451965,0.08211755,0.0008867293,0.00025065074,0.0030671635,0.001623273,0.00043246988,0.056312326],"genre_scores_gemma":[0.90256035,0.00025622032,0.088660434,0.00018701695,0.00004482012,0.0031198845,0.0030293209,0.0000962711,0.0020456812],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9370538,0.019631969,0.006947459,0.0019171658,0.033036616,0.0014130783],"domain_scores_gemma":[0.84297377,0.055727948,0.012448877,0.016962204,0.06970368,0.0021835994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04564538,0.00063887046,0.00054193696,0.0055141277,0.0007761672,0.0019622853,0.0012708391,0.0010615506,0.001569815],"category_scores_gemma":[0.15471706,0.00032619966,0.0010813083,0.004732814,0.0013534513,0.0025824807,0.0020293728,0.0015712021,0.00089643605],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066820026,0.0037173019,0.48685753,0.00080676854,0.00035469103,0.00026560266,0.007209264,0.01611391,0.012891664,0.02777614,0.010627385,0.43271166],"study_design_scores_gemma":[0.00029613715,0.005902258,0.860414,0.0009547352,0.00016763264,0.0004950065,0.0058594104,0.04264056,0.0197738,0.009221892,0.054063704,0.00021091473],"about_ca_topic_score_codex":0.0037871231,"about_ca_topic_score_gemma":0.0036151765,"teacher_disagreement_score":0.04564538,"about_ca_system_score_codex":0.0022320184,"about_ca_system_score_gemma":0.005589456,"threshold_uncertainty_score":0.24139869},"labels":[],"label_agreement":null},{"id":"W2167228371","doi":"10.1109/ares.2012.43","title":"A Comparative Study of Software Security Pattern Classifications","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software security assurance; Selection (genetic algorithm); Classification scheme; Context (archaeology); Security testing; Computer security model; Software; Scheme (mathematics); Security information and event management; Data mining; Security service; Cloud computing security; Computer security; Machine learning; Information security; Mathematics","score_opus":0.06584608257585717,"score_gpt":0.33490771623700105,"score_spread":0.2690616336611439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167228371","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91995806,0.006231393,0.032815333,0.0013793309,0.00013042807,0.00045391606,0.0010365172,0.0003372069,0.037657894],"genre_scores_gemma":[0.9781708,0.0011252948,0.017737536,0.00008807688,0.00003914136,0.00014780066,0.0011156371,0.00006335318,0.0015123793],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98603034,0.0041954545,0.0018595321,0.0009035357,0.0063423375,0.0006687434],"domain_scores_gemma":[0.82213044,0.1214561,0.013046355,0.005324372,0.035624042,0.0024187374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011666056,0.0004292046,0.00060772686,0.016004259,0.001470441,0.0040615974,0.0011155928,0.00082941493,0.0031015554],"category_scores_gemma":[0.08517173,0.00025835793,0.00084519107,0.01271979,0.0013526614,0.008077783,0.0015147252,0.000898744,0.00059483387],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018553869,0.0004250542,0.40650734,0.001869329,0.00031765486,0.00045457372,0.013232914,0.0033095537,0.005448378,0.03027143,0.0054511665,0.5308572],"study_design_scores_gemma":[0.00022056582,0.0027663633,0.7505435,0.0018810489,0.0006627619,0.00303657,0.05654339,0.07329671,0.009082783,0.047205914,0.054522995,0.00023750708],"about_ca_topic_score_codex":0.0027776263,"about_ca_topic_score_gemma":0.0034595795,"teacher_disagreement_score":0.016004259,"about_ca_system_score_codex":0.0024523549,"about_ca_system_score_gemma":0.0015318437,"threshold_uncertainty_score":0.061696768},"labels":[],"label_agreement":null},{"id":"W2167428235","doi":"10.1109/wicsa.2005.8","title":"ACCA: An Architecture-Centric Concern Analysis Method","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Software architecture; Traceability; Computer science; Architecture; Software engineering; Reference architecture; Software; Architectural pattern; Multilayered architecture; Key (lock); Software architecture description; Architecture tradeoff analysis method; Software development; Software design; Programming language; Computer security; Geography","score_opus":0.019324706820463233,"score_gpt":0.31017990023038366,"score_spread":0.29085519340992044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167428235","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004451061,0.00011898044,0.9842928,0.0003402622,0.000060400755,0.00073727564,0.00042818685,0.003506106,0.0060649733],"genre_scores_gemma":[0.049910534,0.000087566055,0.94601357,0.0001383886,0.000036257534,0.0011633754,0.00049260067,0.00036959042,0.0017880674],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99178714,0.0029809703,0.00044596015,0.0012525823,0.003274057,0.00025932843],"domain_scores_gemma":[0.97251475,0.015065453,0.0017427199,0.0038284834,0.0063691963,0.00047933272],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058801402,0.0021709185,0.0008635951,0.010691735,0.0021005685,0.0037746225,0.002628792,0.002401025,0.007800484],"category_scores_gemma":[0.031468324,0.00091286504,0.002378281,0.0050065154,0.0026410888,0.0043992503,0.0029427581,0.0026803075,0.0025336756],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018099684,0.00035503358,0.013924188,0.0011054354,0.00030254043,0.0005934214,0.006266407,0.014568219,0.00872659,0.16810952,0.023199944,0.76266766],"study_design_scores_gemma":[0.00020850336,0.00036628623,0.011316702,0.00065689796,0.0005269257,0.002574435,0.0041498565,0.45892432,0.016170999,0.30407357,0.20065027,0.00038115538],"about_ca_topic_score_codex":0.004989494,"about_ca_topic_score_gemma":0.00794697,"teacher_disagreement_score":0.010691735,"about_ca_system_score_codex":0.0016945662,"about_ca_system_score_gemma":0.006772881,"threshold_uncertainty_score":0.03109759},"labels":[],"label_agreement":null},{"id":"W2167599600","doi":"10.1109/wpc.2004.1311045","title":"Understanding class evolution in object-oriented software","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software evolution; Computer science; Class (philosophy); Unified Modeling Language; Class diagram; Object-oriented design; Software system; Object-oriented programming; Software engineering; Sequence diagram; Taxonomy (biology); Software design; Programming language; Software; Software development; Artificial intelligence; Software construction","score_opus":0.04434548409233375,"score_gpt":0.2659430897624063,"score_spread":0.22159760567007258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167599600","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35397476,0.0013131901,0.6351969,0.0021431718,0.00004715671,0.00019644786,0.00016321319,0.00062656385,0.006338537],"genre_scores_gemma":[0.5869871,0.00073488994,0.4099363,0.00020114171,0.000026357167,0.00014412616,0.0005026968,0.0001564488,0.0013108728],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99723583,0.0011245875,0.00025325562,0.00038320897,0.0008376861,0.00016553656],"domain_scores_gemma":[0.9826134,0.010391741,0.0026643605,0.0019921823,0.0020003761,0.00033798098],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005007197,0.00041822996,0.00050369144,0.005023094,0.0017772577,0.0037081,0.0013381903,0.0020002096,0.00064998964],"category_scores_gemma":[0.028299948,0.0006108223,0.0006590177,0.003519988,0.0021529621,0.012946246,0.0018599594,0.0015945338,0.00020216354],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001706041,0.0003802077,0.10420998,0.0005062971,0.00006712847,0.0011049343,0.036497373,0.037233416,0.008752114,0.32863414,0.0024812394,0.47996256],"study_design_scores_gemma":[0.000041566025,0.00012945276,0.045902055,0.00038441076,0.000085923835,0.0017484479,0.010761422,0.36697525,0.009325868,0.51118416,0.05336051,0.00010090132],"about_ca_topic_score_codex":0.012120694,"about_ca_topic_score_gemma":0.010639899,"teacher_disagreement_score":0.012120694,"about_ca_system_score_codex":0.002608719,"about_ca_system_score_gemma":0.001627869,"threshold_uncertainty_score":0.026480854},"labels":[],"label_agreement":null},{"id":"W2167809408","doi":"10.1109/wcre.1999.806964","title":"Experiments with clustering as a software remodularization method","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":311,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Cluster analysis; Computer science; Reverse engineering; Data mining; Domain (mathematical analysis); Software; Code (set theory); Source code; Software engineering; Theoretical computer science; Machine learning; Programming language; Mathematics","score_opus":0.021145134037425826,"score_gpt":0.3051099445116784,"score_spread":0.28396481047425254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167809408","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9514106,0.00060801336,0.038712382,0.00036742547,0.00015441781,0.0011239477,0.0009852236,0.0025058175,0.004132145],"genre_scores_gemma":[0.8108128,0.00024001852,0.1815551,0.00014423177,0.00005602191,0.0009299564,0.0029931255,0.000497044,0.0027716875],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9924259,0.0038652637,0.0007836439,0.0010892329,0.0013724043,0.0004636014],"domain_scores_gemma":[0.93801427,0.04376283,0.0020043466,0.0074192975,0.007382265,0.001416965],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008327816,0.0014995814,0.0013089946,0.0020912974,0.0015883702,0.0010678057,0.0022351535,0.0019067073,0.0018708888],"category_scores_gemma":[0.039239764,0.0005663837,0.00095048716,0.0032962349,0.0010259032,0.0019621274,0.0013946327,0.0015715415,0.00077196973],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.01409724,0.014587012,0.036397677,0.0027300962,0.0010085407,0.00063005043,0.004250519,0.37826225,0.055302385,0.005682372,0.014984239,0.47206765],"study_design_scores_gemma":[0.0011881592,0.0069008917,0.021139216,0.000090749054,0.00033887898,0.00033909426,0.0015118599,0.90913314,0.048358764,0.0052727265,0.005534264,0.00019220995],"about_ca_topic_score_codex":0.0063902643,"about_ca_topic_score_gemma":0.0055333325,"teacher_disagreement_score":0.008327816,"about_ca_system_score_codex":0.0012028035,"about_ca_system_score_gemma":0.000953148,"threshold_uncertainty_score":0.04404223},"labels":[],"label_agreement":null},{"id":"W2167984494","doi":"10.1145/1218563.1218588","title":"Efficiently mining crosscutting concerns through random walks","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Java; Domain (mathematical analysis); Aspect-oriented programming; Random walk; Data mining; Popularity; Rank (graph theory); AspectJ; Middleware (distributed applications); Theoretical computer science; Database; Programming language; Software; Mathematics; Statistics","score_opus":0.030682819422939444,"score_gpt":0.32147693706841196,"score_spread":0.2907941176454725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167984494","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15412894,0.00030974456,0.84197754,0.000167669,0.000015019313,0.00011253629,0.00029165094,0.0021861321,0.0008107212],"genre_scores_gemma":[0.59249824,0.00022167029,0.40339944,0.00010442133,0.000032034008,0.00020062408,0.0016689624,0.0003427215,0.0015317948],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985892,0.0004193439,0.00009031813,0.00036758464,0.00042035372,0.0001131433],"domain_scores_gemma":[0.9918249,0.00624674,0.0005941523,0.00060877093,0.00058720115,0.00013829226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014997828,0.0009613471,0.0011505252,0.0033538041,0.00055483886,0.0012519045,0.0011851647,0.0010374686,0.0007390273],"category_scores_gemma":[0.0094238175,0.0006659558,0.0011957894,0.0018580977,0.0005223064,0.0017277729,0.00086687464,0.00084181735,0.00046935014],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041988635,0.00033613242,0.049154866,0.000479824,0.00030902913,0.0011418642,0.00074171124,0.5066795,0.025161067,0.024659341,0.0054995674,0.38541728],"study_design_scores_gemma":[0.000014726225,0.000038419887,0.0010702936,0.000010518498,0.00001847752,0.00014463867,0.00004146134,0.9840831,0.0018837078,0.011914994,0.00076998596,0.00000964022],"about_ca_topic_score_codex":0.003074291,"about_ca_topic_score_gemma":0.006139462,"teacher_disagreement_score":0.0033538041,"about_ca_system_score_codex":0.00044522076,"about_ca_system_score_gemma":0.00077279774,"threshold_uncertainty_score":0.007931709},"labels":[],"label_agreement":null},{"id":"W2167986768","doi":"10.1109/csse.2008.1400","title":"Investigation on Academic Research Software Development","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Algoma University","funders":"","keywords":"Software development; Personal software process; Software engineering; Computer science; Software construction; Software peer review; Documentation; Software documentation; Package development process; Software development process; Software quality; Software analytics; Social software engineering; Software project management; Software; Goal-Driven Software Development Process; Programming language","score_opus":0.22323956861653266,"score_gpt":0.348216006006331,"score_spread":0.12497643738979836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167986768","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92373425,0.0014967219,0.0039546913,0.005503216,0.000089932895,0.00049237517,0.0001629613,0.000048586553,0.06451732],"genre_scores_gemma":[0.9917013,0.0010100823,0.0020423515,0.0009509963,0.000082178696,0.0003536805,0.00011136435,0.00003256266,0.0037154052],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.94276255,0.029603807,0.004590212,0.0024234203,0.017441051,0.003178916],"domain_scores_gemma":[0.67053306,0.20032431,0.034857064,0.0109177455,0.07058739,0.0127805015],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.038705595,0.00033070814,0.0006238234,0.011798617,0.006581985,0.0068667396,0.0015430194,0.0010160354,0.004937055],"category_scores_gemma":[0.17602703,0.00040981604,0.00035996235,0.014943413,0.00482316,0.006042594,0.0057836617,0.0019431646,0.0009527612],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020186215,0.000860039,0.36689997,0.0014227267,0.000041382165,0.0010302261,0.37825435,0.00034873828,0.0018793357,0.08125011,0.0063302284,0.16148102],"study_design_scores_gemma":[0.000047959024,0.0006143446,0.27894294,0.001608238,0.000046554997,0.0010115305,0.54377055,0.0019853597,0.003988133,0.013777225,0.15413149,0.000075668184],"about_ca_topic_score_codex":0.003955396,"about_ca_topic_score_gemma":0.0028138233,"teacher_disagreement_score":0.9612944,"about_ca_system_score_codex":0.009525504,"about_ca_system_score_gemma":0.01354,"threshold_uncertainty_score":0.20469725},"labels":[],"label_agreement":null},{"id":"W2167990851","doi":"10.1109/wcre.2003.1287237","title":"Towards the reverse engineering of UML sequence diagrams","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":68,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Sequence diagram; Unified Modeling Language; Computer science; Reverse engineering; UML tool; Class diagram; Sequence (biology); Applications of UML; Communication diagram; Programming language; Software engineering; Software; Chemistry","score_opus":0.03323764457883913,"score_gpt":0.262192049138497,"score_spread":0.22895440455965785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167990851","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036253182,0.00010051311,0.9942041,0.00009036717,0.000021432134,0.00011600741,0.000039356397,0.0014162217,0.00038663865],"genre_scores_gemma":[0.022507753,0.00025694177,0.9753332,0.000064648244,0.000012147114,0.00012282457,0.00036695754,0.00034082276,0.0009946442],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98032606,0.008794145,0.0016272102,0.0019185018,0.006806994,0.0005270866],"domain_scores_gemma":[0.95086575,0.019669544,0.0034667433,0.014686725,0.010884026,0.00042711062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014373815,0.0015904338,0.0010033188,0.0036956435,0.0010977497,0.0042842682,0.0020406963,0.0018786802,0.0018285121],"category_scores_gemma":[0.04481648,0.0015260057,0.0020949924,0.0016142413,0.002448307,0.0037363165,0.0037524372,0.0041894987,0.0014275318],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016399384,0.0003902137,0.006790103,0.00092693826,0.00018190798,0.00063874194,0.003577324,0.062296595,0.040201742,0.12626874,0.0035949065,0.7549688],"study_design_scores_gemma":[0.00015941073,0.0004037958,0.0019525619,0.0007075195,0.00024097956,0.002254597,0.0012544115,0.5007651,0.1782632,0.17140913,0.14239047,0.00019883005],"about_ca_topic_score_codex":0.004410714,"about_ca_topic_score_gemma":0.0063506607,"teacher_disagreement_score":0.014373815,"about_ca_system_score_codex":0.0013334589,"about_ca_system_score_gemma":0.004382942,"threshold_uncertainty_score":0.0760169},"labels":[],"label_agreement":null},{"id":"W2168462212","doi":"10.1145/2382756.2382786","title":"Issue ownership activity in two large software projects","year":2012,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software development; Workload; Software project management; Computer science; Product (mathematics); Software; Software quality; Process management; Risk analysis (engineering); Software engineering; Business; Software construction","score_opus":0.030713618094349967,"score_gpt":0.2943419916532914,"score_spread":0.2636283735589414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2168462212","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.999637,0.000031264608,0.00013190351,0.000020377065,9.639659e-7,0.0000041048784,0.000028719667,0.0000029943049,0.0001426093],"genre_scores_gemma":[0.9993457,0.000024774748,0.00026529803,0.0000042929755,0.0000033789338,0.000009472135,0.00015588825,0.0000032186713,0.00018797383],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9974681,0.00089210883,0.00022496429,0.00036334016,0.0007466153,0.000304797],"domain_scores_gemma":[0.9479036,0.022401178,0.020172063,0.0020457876,0.004135736,0.0033416548],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049605477,0.00024163401,0.0002227024,0.0031492873,0.0006866342,0.0011253667,0.0005961634,0.0004777263,0.00080648693],"category_scores_gemma":[0.030796181,0.0003091151,0.00026749252,0.002712166,0.00092245045,0.0015575049,0.0018320933,0.00075053924,0.00015743896],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019504572,0.00016549339,0.98371744,0.000021736805,0.000027274255,0.00025614252,0.0033391137,0.00064531935,0.00069595955,0.0003473682,0.00017641949,0.010412677],"study_design_scores_gemma":[0.0000058125024,0.00013501791,0.9931825,0.0000067408146,0.0000117806785,0.00017619549,0.0021202443,0.0034118998,0.00027681034,0.0002523085,0.00040792182,0.000012721561],"about_ca_topic_score_codex":0.005079097,"about_ca_topic_score_gemma":0.0064136772,"teacher_disagreement_score":0.005079097,"about_ca_system_score_codex":0.0008506909,"about_ca_system_score_gemma":0.00052973186,"threshold_uncertainty_score":0.02623421},"labels":[],"label_agreement":null},{"id":"W2168482931","doi":"10.1109/ccece.2003.1226137","title":"Metrics for agent-based software development","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Function point; Cohesion (chemistry); Heuristics; Measure (data warehouse); Software engineering; Multi-agent system; Software development; Software metric; Software; Artificial intelligence; Data mining; Software quality; Programming language","score_opus":0.03767505704648115,"score_gpt":0.27960903538026016,"score_spread":0.241933978333779,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2168482931","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051630074,0.0065565426,0.97631854,0.0008982184,0.00029962367,0.00055503275,0.00053477864,0.0011125789,0.008561723],"genre_scores_gemma":[0.11055829,0.0036062228,0.8799883,0.00014943114,0.00018813208,0.0017395318,0.0015756133,0.0003826006,0.0018118687],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9764138,0.008872349,0.0031305884,0.0013294385,0.00992782,0.00032601296],"domain_scores_gemma":[0.96268094,0.017757503,0.0050736438,0.003875623,0.009607777,0.0010045216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013224198,0.0025254558,0.0015052553,0.008049479,0.0011701611,0.004970429,0.0018943379,0.0017070514,0.0022438702],"category_scores_gemma":[0.07606523,0.00059367105,0.0009170207,0.009853164,0.0019068815,0.0061464463,0.0034162335,0.0022320917,0.0009584931],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008170098,0.00013086344,0.0054436736,0.0016526581,0.00018246954,0.0001705817,0.00069537846,0.09831177,0.0030573697,0.47653404,0.011915786,0.40182367],"study_design_scores_gemma":[0.000043340307,0.00044530252,0.0044100494,0.0010404466,0.00010803724,0.00043135966,0.00042980397,0.2258433,0.0042927167,0.63388634,0.12887217,0.00019713462],"about_ca_topic_score_codex":0.002401956,"about_ca_topic_score_gemma":0.001498854,"teacher_disagreement_score":0.013224198,"about_ca_system_score_codex":0.0032316365,"about_ca_system_score_gemma":0.002515429,"threshold_uncertainty_score":0.06993711},"labels":[],"label_agreement":null},{"id":"W2168499709","doi":"10.1145/1569901.1570125","title":"Software project planning for robustness and completion time in the presence of uncertainty using multi objective search based software engineering","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Robustness (evolution); Search-based software engineering; Computer science; Software sizing; Software project management; Software; Software construction; Software development; Software metric; Reliability engineering; Software engineering; Engineering","score_opus":0.054123703457164406,"score_gpt":0.31675709318542333,"score_spread":0.26263338972825895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2168499709","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20796794,0.00055142323,0.78795105,0.0004805008,0.0000137380775,0.000098772965,0.000073603005,0.0001679438,0.0026951171],"genre_scores_gemma":[0.91331434,0.00015758636,0.08559852,0.000021862441,0.0000094936995,0.00012697997,0.0000592268,0.000047648835,0.0006643387],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978436,0.0014464934,0.000069878864,0.00015891084,0.00035430896,0.00012674529],"domain_scores_gemma":[0.9885302,0.010024816,0.00066741917,0.00016702403,0.0004173753,0.00019309345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004644374,0.0007659343,0.0012968705,0.0013423555,0.0003806335,0.0013150033,0.0006737669,0.00088733097,0.0012929006],"category_scores_gemma":[0.013790825,0.0006523849,0.00094070676,0.0012609893,0.0009097877,0.0016915455,0.0009102395,0.00093658577,0.00008676334],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039818144,0.000015150898,0.00022640199,0.000021031328,0.000019207953,0.000015590793,0.000024112953,0.9926927,0.00027604445,0.0025632891,0.00006470671,0.004041936],"study_design_scores_gemma":[0.0000071101545,0.000021722137,0.00014311912,0.000003357933,0.0000049223927,0.0000038130643,0.000008730456,0.9977254,0.00016223988,0.0018697815,0.000045764347,0.0000040422938],"about_ca_topic_score_codex":0.008170333,"about_ca_topic_score_gemma":0.004309811,"teacher_disagreement_score":0.008170333,"about_ca_system_score_codex":0.0016790967,"about_ca_system_score_gemma":0.0023168405,"threshold_uncertainty_score":0.02456212},"labels":[],"label_agreement":null},{"id":"W2168709131","doi":"10.1109/wcre.1995.514708","title":"Formal representation of reuseable software modules","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Executable; Programming language; Notation; Algebraic specification; Theoretical computer science; Formalism (music); Software; Formal specification; Mathematics","score_opus":0.036457099871575256,"score_gpt":0.26839960245578554,"score_spread":0.23194250258421029,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2168709131","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009610954,0.00037402593,0.98029584,0.00041730146,0.000053462278,0.0001152421,0.00024384308,0.00060938834,0.008279914],"genre_scores_gemma":[0.25889155,0.0009347292,0.7275872,0.0002313211,0.00014865327,0.00058543467,0.0011003643,0.00021952104,0.010301273],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9981687,0.00048979017,0.0002060903,0.00019798958,0.00073681335,0.000200588],"domain_scores_gemma":[0.99776006,0.00075009046,0.0004365565,0.0004708663,0.00048412717,0.00009833192],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021735716,0.00065615796,0.00038648688,0.0017938074,0.00068090175,0.0027638713,0.0019272134,0.001336278,0.0027865507],"category_scores_gemma":[0.0036748655,0.00059977703,0.0014136557,0.0012325452,0.0030341577,0.003308568,0.0013883454,0.0014712933,0.0008889676],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011750866,0.000029371522,0.00012124532,0.00007607814,0.000012432238,0.0001889527,0.00035480232,0.012890913,0.0025989101,0.972305,0.00058872154,0.010821881],"study_design_scores_gemma":[0.000042508495,0.00006268128,0.00015934855,0.00010412481,0.000037947633,0.00037170228,0.00013831575,0.09573392,0.0048607634,0.8464611,0.05199166,0.000035949724],"about_ca_topic_score_codex":0.0020052332,"about_ca_topic_score_gemma":0.002165315,"teacher_disagreement_score":0.0027865507,"about_ca_system_score_codex":0.0014346082,"about_ca_system_score_gemma":0.0016074453,"threshold_uncertainty_score":0.011495113},"labels":[],"label_agreement":null},{"id":"W2168871565","doi":"10.1145/1297846.1297934","title":"Mining implementation recipes of framework-provided concepts in dynamic framework API interaction traces","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Eclipse; Computer science; Program comprehension; Slicing; Context (archaeology); Software engineering; Code (set theory); Programming language; Debugging; Plug-in; Program slicing; World Wide Web; Software; Software system","score_opus":0.02245870803812806,"score_gpt":0.39207922770022924,"score_spread":0.3696205196621012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2168871565","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22738624,0.0003009467,0.7566937,0.00020099443,0.000049554805,0.00036728868,0.002231011,0.01025495,0.0025152955],"genre_scores_gemma":[0.2838328,0.00020227562,0.70798796,0.00003916306,0.000010679949,0.00032854907,0.0036014265,0.0018822066,0.0021149104],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99882,0.00018781888,0.00011399159,0.00036388604,0.00041683906,0.00009748571],"domain_scores_gemma":[0.992224,0.004194426,0.00079566013,0.0012040038,0.0013638401,0.00021807794],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018316235,0.000829228,0.00043757271,0.0028586385,0.00065280375,0.0015392276,0.0013772742,0.0008441073,0.0014542125],"category_scores_gemma":[0.016097186,0.0009182282,0.00093630113,0.0015526947,0.00074257614,0.0025097737,0.0010119611,0.001361763,0.00049101113],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066357007,0.00041899722,0.09863232,0.0013856194,0.00016292703,0.0037774416,0.009828634,0.062921286,0.07785802,0.06976543,0.012844077,0.6617417],"study_design_scores_gemma":[0.000073815296,0.00022156282,0.03195594,0.00033408834,0.00017442986,0.001543312,0.0025266723,0.7206551,0.11990746,0.052833848,0.06950378,0.0002701173],"about_ca_topic_score_codex":0.006960031,"about_ca_topic_score_gemma":0.012361695,"teacher_disagreement_score":0.006960031,"about_ca_system_score_codex":0.0010144762,"about_ca_system_score_gemma":0.001938382,"threshold_uncertainty_score":0.013839066},"labels":[],"label_agreement":null},{"id":"W2168972892","doi":"10.1023/a:1024460315173","title":"User Interface Reverse Engineering in Support of Interface Migration to the Web","year":2003,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Legacy system; Interface (matter); User interface; Context (archaeology); TRACE (psycholinguistics); Process (computing); Human–computer interaction; State (computer science); Task (project management); Interface metaphor; Transition (genetics); Source code; User interface design; World Wide Web; Graphical user interface testing; Programming language; Operating system; Software; Engineering","score_opus":0.008918239222152284,"score_gpt":0.24989148076236092,"score_spread":0.24097324154020863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2168972892","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3206641,0.00056611205,0.5918149,0.0011703671,0.0004516276,0.00035052496,0.00016712757,0.07106647,0.013748792],"genre_scores_gemma":[0.8109363,0.0002094113,0.17391652,0.00045813547,0.000092810216,0.000089762274,0.00028203407,0.0021960794,0.011818883],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979678,0.0007014568,0.00016664938,0.00024376206,0.0006683126,0.00025192718],"domain_scores_gemma":[0.9859886,0.003998636,0.00076562504,0.0062742042,0.0026778402,0.00029495356],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025203342,0.00060476124,0.0005821674,0.00077212387,0.0006918315,0.0018227657,0.0017810159,0.0014394913,0.0034873695],"category_scores_gemma":[0.018100988,0.0005015339,0.0004927581,0.00044738725,0.0005690573,0.0022013204,0.0015387316,0.0018316443,0.0013225173],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024742514,0.001147178,0.017655684,0.0004232239,0.00014362273,0.0020518566,0.0020752987,0.009470856,0.12385004,0.011953562,0.017077636,0.8116768],"study_design_scores_gemma":[0.0003590607,0.0006881685,0.010502603,0.00014106592,0.00036908395,0.0029934016,0.00070814055,0.45638084,0.46889368,0.015820574,0.04300142,0.00014198398],"about_ca_topic_score_codex":0.002111797,"about_ca_topic_score_gemma":0.002617025,"teacher_disagreement_score":0.0034873695,"about_ca_system_score_codex":0.0003282912,"about_ca_system_score_gemma":0.0010812053,"threshold_uncertainty_score":0.0133289695},"labels":[],"label_agreement":null},{"id":"W2169584729","doi":"10.1109/euromicro.2007.53","title":"Scope Management of Non-Functional Requirements","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Scope (computer science); Non-functional requirement; Requirements management; Functional requirement; Risk analysis (engineering); Requirement prioritization; Software requirements; Software project management; Software engineering; Process (computing); Requirements analysis; Software development; Software; Systems engineering; Process management; Engineering; Software construction; Business","score_opus":0.032547410864147204,"score_gpt":0.3023089711822597,"score_spread":0.2697615603181125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169584729","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06630079,0.0014550801,0.9143962,0.00052555965,0.00005668784,0.0005027836,0.0001973399,0.0014660099,0.015099526],"genre_scores_gemma":[0.5859294,0.0010924236,0.4063236,0.00012834392,0.00008681923,0.00073442253,0.0007744929,0.00040081458,0.0045297495],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9892777,0.004037049,0.0010289772,0.0009717149,0.004350989,0.0003335931],"domain_scores_gemma":[0.96487033,0.017515844,0.0036793598,0.0052669924,0.0074832644,0.0011841395],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010288257,0.00080902624,0.00058098085,0.0049626455,0.0009363573,0.0024849547,0.0016992424,0.0007953433,0.0019577902],"category_scores_gemma":[0.03902015,0.00053562253,0.0006169234,0.0018691719,0.0010043082,0.0043278844,0.0037589318,0.0009258237,0.0005300819],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020947172,0.00016886673,0.010521561,0.00088952837,0.00007965153,0.00059432676,0.006550388,0.029013148,0.034888398,0.086344786,0.0032125325,0.82752734],"study_design_scores_gemma":[0.00012457542,0.0011969128,0.049091432,0.0014840309,0.00028412556,0.002577948,0.0077462746,0.44992515,0.07347653,0.2548884,0.1588322,0.00037236072],"about_ca_topic_score_codex":0.0013758726,"about_ca_topic_score_gemma":0.0014991888,"teacher_disagreement_score":0.010288257,"about_ca_system_score_codex":0.0008346026,"about_ca_system_score_gemma":0.001499627,"threshold_uncertainty_score":0.05441016},"labels":[],"label_agreement":null},{"id":"W2169719505","doi":"10.1109/icsme.2014.26","title":"Why Do Automated Builds Break? An Empirical Study","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":74,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Leverage (statistics); Executable; Computer science; Software engineering; Breakage; Software; Empirical research; Work (physics); Engineering; World Wide Web; Operating system; Artificial intelligence","score_opus":0.02278950590193803,"score_gpt":0.3377818139732754,"score_spread":0.3149923080713374,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169719505","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99729365,0.00022289145,0.0006035341,0.00028605614,0.00000777487,0.00007750119,0.00017482109,0.000024782797,0.001309033],"genre_scores_gemma":[0.9982804,0.00023480189,0.00070533645,0.0000915057,0.000016380242,0.00006834781,0.00022196489,0.000022636854,0.00035852814],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9812899,0.008202381,0.0018997226,0.0019712306,0.00512104,0.001515744],"domain_scores_gemma":[0.54513156,0.3295613,0.085804746,0.0137582915,0.02104585,0.004698182],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014672207,0.0005769539,0.00045820998,0.005072307,0.001992831,0.0032027392,0.0024744645,0.001882616,0.00313218],"category_scores_gemma":[0.13818568,0.0010506355,0.0004195705,0.0050891526,0.0022600773,0.0049536987,0.0019116557,0.0026198744,0.000825829],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026756478,0.001290457,0.93031466,0.0002968954,0.00010792289,0.0015640603,0.028876081,0.0009806668,0.000577242,0.001026531,0.0021600618,0.032537878],"study_design_scores_gemma":[0.000041143525,0.00065776234,0.92313933,0.0002565085,0.00008302505,0.002056275,0.057202373,0.005334937,0.00081740256,0.0009308361,0.009411324,0.0000690145],"about_ca_topic_score_codex":0.0041899723,"about_ca_topic_score_gemma":0.0062179174,"teacher_disagreement_score":0.9853278,"about_ca_system_score_codex":0.0018110988,"about_ca_system_score_gemma":0.00168064,"threshold_uncertainty_score":0.077594936},"labels":[],"label_agreement":null},{"id":"W2169751269","doi":"10.1145/2702123.2702426","title":"STRATOS","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Alberta Innovates - Technology Futures","keywords":"Plan (archaeology); Computer science; Visualization; Process (computing); Software; Product (mathematics); Process management; Order (exchange); Software engineering; Risk analysis (engineering); Management science; Operations research; Engineering; Business; Data mining","score_opus":0.05014632302939172,"score_gpt":0.28880324012383796,"score_spread":0.23865691709444625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169751269","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023862105,0.0018073492,0.50591946,0.0023926103,0.0005031553,0.000992906,0.025066601,0.2855412,0.15391463],"genre_scores_gemma":[0.24172455,0.0030021628,0.5786791,0.00135605,0.00016591397,0.0020954867,0.03593676,0.05113757,0.08590234],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992059,0.00024932914,0.00006297994,0.00015420694,0.000233981,0.000093689516],"domain_scores_gemma":[0.99649173,0.0021965173,0.0002272901,0.0005194363,0.00034095792,0.00022409686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015779323,0.001466233,0.0005349893,0.001967389,0.0008265463,0.0028510334,0.001635993,0.0012211645,0.07626123],"category_scores_gemma":[0.0072188,0.00086039514,0.0011036989,0.001022986,0.00086456374,0.003852248,0.0028277214,0.0013382005,0.019305566],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015718893,0.00030228766,0.0045590084,0.00326062,0.0001260147,0.0011914482,0.004944884,0.01565911,0.016093181,0.15810272,0.40566924,0.3885196],"study_design_scores_gemma":[0.00028916603,0.00026640063,0.0021376843,0.0007799296,0.000067964,0.00084138516,0.0011157305,0.034664635,0.00938156,0.08145195,0.86887187,0.00013172967],"about_ca_topic_score_codex":0.0025974766,"about_ca_topic_score_gemma":0.006062071,"teacher_disagreement_score":0.07626123,"about_ca_system_score_codex":0.000872043,"about_ca_system_score_gemma":0.0018271453,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2169782617","doi":"10.1109/tse.2002.1158285","title":"An operational process for goal-driven definition of measures","year":2002,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":148,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Empirical research; Set (abstract data type); Process (computing); Verifiable secret sharing; Software engineering; Management science; Programming language","score_opus":0.03966221107474621,"score_gpt":0.26917200462530516,"score_spread":0.22950979355055895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169782617","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00072710303,0.00017949686,0.99262685,0.0016336176,0.000100166726,0.0004068279,0.000112856906,0.00016638769,0.004046602],"genre_scores_gemma":[0.026703719,0.00021540969,0.9690554,0.0006048627,0.000101100966,0.0022152371,0.00030198204,0.00012063651,0.0006816884],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96719915,0.018368095,0.003739479,0.0035811393,0.006305089,0.0008069801],"domain_scores_gemma":[0.9493905,0.025320232,0.0032466175,0.009038229,0.011549419,0.001455112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.051249042,0.0026093507,0.0022272915,0.007595867,0.0032181372,0.009940111,0.006977259,0.004856401,0.005317696],"category_scores_gemma":[0.07780371,0.0014430253,0.0033917848,0.0062034084,0.01527843,0.01661984,0.00835169,0.010999686,0.002701184],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008925061,0.000033803568,0.000112770365,0.00010471282,0.000014464522,0.000056614743,0.0010421758,0.0010929777,0.00033072018,0.986704,0.00085105846,0.009647859],"study_design_scores_gemma":[0.000027672078,0.00007427421,0.00019903842,0.00025554287,0.000020697205,0.00012401825,0.00067452993,0.011137103,0.0008528931,0.9493262,0.037257023,0.00005110199],"about_ca_topic_score_codex":0.0024456216,"about_ca_topic_score_gemma":0.0014111458,"teacher_disagreement_score":0.051249042,"about_ca_system_score_codex":0.0041987393,"about_ca_system_score_gemma":0.008133901,"threshold_uncertainty_score":0.27103406},"labels":[],"label_agreement":null},{"id":"W2169917206","doi":"10.1109/msr.2007.4","title":"Correlating Social Interactions to Release History during Software Evolution","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software evolution; Set (abstract data type); Source code; Code (set theory); Software; Similarity (geometry); Minor (academic); Information retrieval; Exploratory research; Data mining; Software system; Artificial intelligence; Programming language; Software construction","score_opus":0.018640779547239392,"score_gpt":0.2687361363165684,"score_spread":0.250095356769329,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169917206","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90304166,0.0011492223,0.088619,0.0005886195,0.000064791784,0.00024984425,0.0015154508,0.00085539155,0.0039160196],"genre_scores_gemma":[0.95917416,0.0003657954,0.037499197,0.00006742572,0.00012379496,0.00013784549,0.0014766656,0.000105103005,0.0010498666],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9961062,0.001462203,0.0003922539,0.00074816257,0.0010915109,0.00019973758],"domain_scores_gemma":[0.9416685,0.037183,0.012457198,0.0033102122,0.004437858,0.0009432147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040443167,0.00057507167,0.000666756,0.010456754,0.0009132761,0.002242279,0.00096683676,0.001102795,0.001465214],"category_scores_gemma":[0.032526843,0.0005705813,0.00061354437,0.0057443934,0.0006939103,0.003601604,0.0012474813,0.0009440644,0.00055770675],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004633879,0.00036931402,0.7901609,0.0004650592,0.00035631828,0.0006091224,0.0046357913,0.006215979,0.009734389,0.0017048838,0.0012300862,0.18405472],"study_design_scores_gemma":[0.000049051676,0.0005196683,0.8671197,0.00011853814,0.00030793934,0.0010130348,0.0035210666,0.10185703,0.010591066,0.007190978,0.0075114653,0.00020041014],"about_ca_topic_score_codex":0.0046052663,"about_ca_topic_score_gemma":0.0058242353,"teacher_disagreement_score":0.010456754,"about_ca_system_score_codex":0.00063355005,"about_ca_system_score_gemma":0.00080618233,"threshold_uncertainty_score":0.02138859},"labels":[],"label_agreement":null},{"id":"W2169950128","doi":"10.1109/wcre.2000.891477","title":"ACCD: an algorithm for comprehension-driven clustering","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":252,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of Waterloo","funders":"","keywords":"Program comprehension; Computer science; Cluster analysis; Cohesion (chemistry); Comprehension; Software; Data mining; Software maintenance; Software engineering; Algorithm; Theoretical computer science; Software system; Programming language; Artificial intelligence","score_opus":0.05190595391670877,"score_gpt":0.2896389560241501,"score_spread":0.2377330021074413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169950128","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015650142,0.00013627487,0.98492706,0.00015569314,0.00011569382,0.00036378283,0.0005240834,0.0111248465,0.0010875018],"genre_scores_gemma":[0.007553161,0.00005646376,0.98726857,0.00013266664,0.00005941372,0.0005594464,0.0016880605,0.001139717,0.001542551],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99582833,0.0009869394,0.00039471535,0.0010073797,0.001514179,0.00026850667],"domain_scores_gemma":[0.9922563,0.0027035244,0.00033584944,0.0013132946,0.0031051638,0.00028594327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042211283,0.003432658,0.0021729057,0.00908733,0.0033758776,0.003418369,0.006291768,0.004889153,0.016204437],"category_scores_gemma":[0.01887752,0.0016640575,0.0025887487,0.0071534324,0.0016182446,0.0042037508,0.004875307,0.0038668555,0.015325074],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035034606,0.00019431148,0.0018553954,0.00042991276,0.0001777269,0.00013853508,0.00038503553,0.039106917,0.0064689787,0.012007853,0.06895627,0.86992866],"study_design_scores_gemma":[0.00037495914,0.00014738733,0.0010817114,0.00010470082,0.000105999774,0.0005730583,0.00025957334,0.8600719,0.013228433,0.051763587,0.07214869,0.00013994641],"about_ca_topic_score_codex":0.007137279,"about_ca_topic_score_gemma":0.012845329,"teacher_disagreement_score":0.016204437,"about_ca_system_score_codex":0.0019852077,"about_ca_system_score_gemma":0.0040763333,"threshold_uncertainty_score":0.054209292},"labels":[],"label_agreement":null},{"id":"W2169952536","doi":"10.1145/1062455.1062491","title":"Using structural context to recommend source code examples","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":385,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Programming language; Source code; Eclipse; Coding (social sciences); Code (set theory); Context (archaeology); Code review; Class (philosophy); Code generation; Task (project management); Software engineering; Static program analysis; Artificial intelligence; Software development; Software","score_opus":0.0811007841125726,"score_gpt":0.32975704602887224,"score_spread":0.24865626191629964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169952536","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5276095,0.002877956,0.44372022,0.0012368289,0.00015518931,0.0009533717,0.0036248616,0.009587246,0.010234826],"genre_scores_gemma":[0.6395274,0.00077158824,0.34980452,0.00016248721,0.000071284565,0.00041767335,0.0068509653,0.0004280914,0.0019660164],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99862814,0.00034583217,0.00010934112,0.00035609026,0.00049038633,0.00007034921],"domain_scores_gemma":[0.98686916,0.008383459,0.00082150503,0.00087880023,0.0027423108,0.00030483084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014275985,0.0008280496,0.0007146724,0.010199413,0.000843219,0.0013348483,0.0011811913,0.0014932617,0.00242333],"category_scores_gemma":[0.020630866,0.000580225,0.0005437072,0.0037969635,0.00040003465,0.002280326,0.00090143253,0.00091307255,0.0011186039],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066100573,0.0007234836,0.119832136,0.0011382222,0.00023932107,0.0009789993,0.0023703277,0.020437438,0.014997514,0.004280449,0.023697976,0.8106432],"study_design_scores_gemma":[0.00040124758,0.0007031921,0.0711083,0.00083348394,0.0007133733,0.002283355,0.0028553254,0.80671185,0.030505056,0.0232048,0.060414243,0.0002658314],"about_ca_topic_score_codex":0.009028036,"about_ca_topic_score_gemma":0.033134278,"teacher_disagreement_score":0.010199413,"about_ca_system_score_codex":0.00059549615,"about_ca_system_score_gemma":0.0015716517,"threshold_uncertainty_score":0.017951012},"labels":[],"label_agreement":null},{"id":"W2169962133","doi":"10.1109/wcre.2002.1173066","title":"A methodology for developing transformations using the maintainability soft-goal graph","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Maintainability; Computer science; Software engineering; Code refactoring; Object-oriented programming; Restructuring; Legacy system; Graph rewriting; Process (computing); Context (archaeology); Graph; Systems engineering; Programming language; Engineering; Theoretical computer science; Software","score_opus":0.12101834387797614,"score_gpt":0.36841855839530635,"score_spread":0.2474002145173302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169962133","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006739726,0.000026400608,0.99731845,0.00009721355,0.000013288282,0.00015801936,0.000034443787,0.00067219895,0.0010060125],"genre_scores_gemma":[0.008072526,0.00007896535,0.99033713,0.00004428305,0.0000072631806,0.00021542898,0.00013334639,0.00024349523,0.00086748815],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99688643,0.0008707788,0.00029679414,0.0006575685,0.0011307787,0.00015777175],"domain_scores_gemma":[0.99480116,0.00227832,0.00053043425,0.001344564,0.00086417893,0.00018129365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00440377,0.0015992983,0.00078987505,0.0039927275,0.0016621004,0.0032789598,0.0025528593,0.0012638723,0.0035248485],"category_scores_gemma":[0.009312252,0.0013880488,0.0029592249,0.0024446761,0.0034333465,0.0035211241,0.002910416,0.0037971642,0.0015168166],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000053074054,0.00025054193,0.0013735826,0.0007467797,0.00012955494,0.0008649101,0.0027742027,0.046238277,0.013505357,0.49480802,0.0059130094,0.43334264],"study_design_scores_gemma":[0.00008806869,0.0003092804,0.0011029763,0.0005143318,0.00026271425,0.0020710884,0.0011159283,0.28771245,0.03251123,0.4869009,0.1871675,0.00024363528],"about_ca_topic_score_codex":0.0048399586,"about_ca_topic_score_gemma":0.007361716,"teacher_disagreement_score":0.0048399586,"about_ca_system_score_codex":0.001620788,"about_ca_system_score_gemma":0.0038770426,"threshold_uncertainty_score":0.02328962},"labels":[],"label_agreement":null},{"id":"W2170146914","doi":"10.1109/ccece.1998.685594","title":"On measuring programmer team productivity","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Programmer; Productivity; Computer science; Context (archaeology); Process (computing); Variance (accounting); Software engineering; Programming language; Accounting; Business","score_opus":0.04951380512404416,"score_gpt":0.25028091670451674,"score_spread":0.20076711158047258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170146914","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69214946,0.008039761,0.23383091,0.0019323326,0.00059219584,0.0010924137,0.001679582,0.0006553935,0.060027976],"genre_scores_gemma":[0.8608962,0.0049715247,0.12412747,0.0008323097,0.0005349436,0.0017273759,0.0018233444,0.00020910091,0.0048776036],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95975995,0.017530326,0.0021447286,0.0031447934,0.016336685,0.0010834545],"domain_scores_gemma":[0.8276753,0.11861936,0.016708352,0.011718193,0.021647856,0.0036309112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025743308,0.0011043499,0.0009775846,0.010214964,0.0016258468,0.0033640696,0.0014640989,0.0017461834,0.0026807194],"category_scores_gemma":[0.11866787,0.00046359637,0.00065815623,0.010629073,0.0020616257,0.0066231685,0.0036872623,0.0014742356,0.0012374299],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007051239,0.0009770918,0.33423442,0.0016124585,0.0005130879,0.00015726002,0.0076442286,0.008565683,0.005730739,0.023699293,0.0076685506,0.608492],"study_design_scores_gemma":[0.0002457164,0.006158833,0.8419425,0.0020137785,0.00043404268,0.0011752875,0.009307435,0.02413373,0.01194585,0.060723856,0.04158179,0.00033712044],"about_ca_topic_score_codex":0.0020256697,"about_ca_topic_score_gemma":0.0018012801,"teacher_disagreement_score":0.025743308,"about_ca_system_score_codex":0.0016364853,"about_ca_system_score_gemma":0.0015953993,"threshold_uncertainty_score":0.13614523},"labels":[],"label_agreement":null},{"id":"W2170347996","doi":"10.1109/cseet.2006.42","title":"Will Johnny/Joanie Make a Good Software Engineer? Are Course Grades Showing the Whole Picture?","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University; National Aeronautics and Space Administration","keywords":"Traceability; Grading (engineering); Software engineering; Computer science; Argument (complex analysis); Course (navigation); Software; Software quality; Quality (philosophy); Software evolution; Software development; Engineering; Software construction; Programming language","score_opus":0.011133895997029237,"score_gpt":0.22978632695132478,"score_spread":0.21865243095429554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170347996","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88792306,0.0009889313,0.0021214793,0.070712194,0.00073237997,0.000030259245,0.0001556084,0.00008536514,0.03725069],"genre_scores_gemma":[0.984953,0.00076450215,0.0016921009,0.0036454157,0.00015643564,0.000020404463,0.00009390941,0.000051141873,0.0086231455],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99834645,0.0004907635,0.00008540741,0.00021903135,0.0005820792,0.00027622934],"domain_scores_gemma":[0.9692816,0.007325031,0.0060859467,0.0009989111,0.00679142,0.009517045],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0037234912,0.00024887943,0.0004140336,0.0011666667,0.0017669441,0.0029848206,0.0004350338,0.001309746,0.006763867],"category_scores_gemma":[0.042681653,0.00022140583,0.00020422337,0.0007507127,0.0017091134,0.0032305508,0.001048292,0.0016424935,0.0027343233],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018448911,0.00057310757,0.8012548,0.00014209074,0.000060900926,0.0007979474,0.0074983574,0.00026203168,0.0014491873,0.0069167847,0.06136579,0.11949456],"study_design_scores_gemma":[0.00003632679,0.0003071155,0.8889866,0.00024844232,0.00006864636,0.0013950292,0.025773568,0.0015011422,0.0022836712,0.01722876,0.062042896,0.00012777997],"about_ca_topic_score_codex":0.005983934,"about_ca_topic_score_gemma":0.020259602,"teacher_disagreement_score":0.9962765,"about_ca_system_score_codex":0.0011839134,"about_ca_system_score_gemma":0.0016265212,"threshold_uncertainty_score":0.022627413},"labels":[],"label_agreement":null},{"id":"W2170365891","doi":"10.1145/2393596.2393654","title":"Do crosscutting concerns cause modularity problems?","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Modularity (biology); Computer science; Comprehension; Process (computing); Program comprehension; Period (music); Risk analysis (engineering); Data science; Software; Software system; Programming language; Physics; Business","score_opus":0.061676067535754546,"score_gpt":0.3240856276900568,"score_spread":0.2624095601543022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170365891","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9785015,0.0011143903,0.011799198,0.002097634,0.000051159393,0.00007711405,0.00022722497,0.00034987516,0.005781977],"genre_scores_gemma":[0.9956624,0.0001807751,0.0031497069,0.00016277605,0.000046875786,0.000022555367,0.00016387882,0.0000779728,0.000533147],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9852841,0.004600651,0.001176883,0.0025899142,0.0052973903,0.0010511719],"domain_scores_gemma":[0.6239284,0.22133337,0.10607496,0.025829038,0.018008633,0.004825523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0135802915,0.0004892507,0.00058674806,0.005224292,0.0012007591,0.0020436407,0.0011808437,0.0014980328,0.0041857567],"category_scores_gemma":[0.18283927,0.0006563647,0.0007136042,0.004167683,0.0026808926,0.005558134,0.0020200904,0.0015475699,0.00053799304],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031152973,0.00020352771,0.86627495,0.00043776978,0.00025421416,0.00097015355,0.0068240217,0.0013027133,0.0038910701,0.005102275,0.0025273343,0.11190037],"study_design_scores_gemma":[0.000041992516,0.00034564064,0.95884347,0.00017355688,0.00017872835,0.003314502,0.0046345345,0.007429582,0.003965175,0.0130062755,0.00801267,0.000053909585],"about_ca_topic_score_codex":0.0025044077,"about_ca_topic_score_gemma":0.0032424678,"teacher_disagreement_score":0.0135802915,"about_ca_system_score_codex":0.0011369901,"about_ca_system_score_gemma":0.0010548,"threshold_uncertainty_score":0.07182032},"labels":[],"label_agreement":null},{"id":"W2170369178","doi":"10.1145/1923947.1923951","title":"Improving program navigation with an active help system","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software engineering; Context (archaeology); Software; Program code; Work (physics); Code (set theory); Human–computer interaction; World Wide Web; Multimedia; Programming language; Engineering","score_opus":0.008570555688205844,"score_gpt":0.26166675277680007,"score_spread":0.2530961970885942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170369178","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7689893,0.00036721933,0.19179577,0.00054398115,0.00008517651,0.00074695295,0.00020902455,0.03395406,0.00330846],"genre_scores_gemma":[0.5192466,0.00022394821,0.47333404,0.00033241665,0.000047288973,0.0004077745,0.0005661118,0.0007614911,0.0050802976],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9980806,0.00066784373,0.00012659379,0.0004896771,0.00049619237,0.00013911435],"domain_scores_gemma":[0.98438346,0.0087009985,0.0016383202,0.0022438136,0.0018508799,0.0011824247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025299138,0.0013763201,0.0006183382,0.0013732368,0.0004968087,0.0013625083,0.0022132979,0.0012856857,0.0036729241],"category_scores_gemma":[0.015453488,0.0006949824,0.00040495236,0.0005376729,0.0004722254,0.0022133747,0.0021664712,0.0010058493,0.001343621],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014489816,0.005431901,0.030073084,0.0013297895,0.00014573902,0.0008114008,0.00756868,0.005242641,0.15583922,0.0012943792,0.014049778,0.77676433],"study_design_scores_gemma":[0.004205966,0.03200958,0.236626,0.0012510477,0.0020649722,0.006524725,0.006482122,0.22556841,0.290762,0.0065534734,0.18665196,0.001299774],"about_ca_topic_score_codex":0.0017223713,"about_ca_topic_score_gemma":0.003186637,"teacher_disagreement_score":0.0036729241,"about_ca_system_score_codex":0.00034445527,"about_ca_system_score_gemma":0.0012281939,"threshold_uncertainty_score":0.013379633},"labels":[],"label_agreement":null},{"id":"W2170491347","doi":"10.1109/esem.2007.10","title":"Impact Analysis of Missing Values on the Prediction Accuracy of Analogy-based Software Effort Estimation Method AQUA","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Missing data; Analogy; Intuition; Computer science; Dependency (UML); Statistics; Quadratic equation; Data mining; Mathematics; Context (archaeology); Algorithm; Artificial intelligence","score_opus":0.029467486228604543,"score_gpt":0.36471030131969917,"score_spread":0.3352428150910946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170491347","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9369316,0.0004724962,0.06004286,0.0003948296,0.00005044744,0.00008454974,0.00030726366,0.0006452003,0.001070807],"genre_scores_gemma":[0.9853755,0.000056902023,0.014027378,0.000034349432,0.000016978533,0.00005296063,0.0002956212,0.000037093436,0.00010325482],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.95541,0.027782138,0.0028461057,0.0031423492,0.009927795,0.0008916361],"domain_scores_gemma":[0.41658527,0.52354515,0.015785135,0.030934066,0.011893317,0.0012571003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04140347,0.0008114441,0.0014334626,0.0020617957,0.00077477057,0.0013797645,0.0014126766,0.0011279371,0.0011578712],"category_scores_gemma":[0.278738,0.0004903011,0.0010794901,0.0017889914,0.0012057694,0.0035092183,0.002060943,0.0017661637,0.00023469054],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0123627735,0.00093441276,0.39648214,0.00081563165,0.0011541499,0.0005895054,0.0017875383,0.24679817,0.009771007,0.0034663104,0.0017772652,0.32406113],"study_design_scores_gemma":[0.0002264249,0.004006204,0.1665637,0.0001987743,0.0005006504,0.0007116658,0.0010296636,0.79114795,0.026227698,0.00764662,0.001577047,0.00016360088],"about_ca_topic_score_codex":0.0015909495,"about_ca_topic_score_gemma":0.0009462296,"teacher_disagreement_score":0.04140347,"about_ca_system_score_codex":0.0008718544,"about_ca_system_score_gemma":0.0008855838,"threshold_uncertainty_score":0.21896511},"labels":[],"label_agreement":null},{"id":"W2170660163","doi":"10.1109/wcre.2010.33","title":"Enhancing Source-Based Clone Detection Using Intermediate Representation","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Source code; Computer science; Software maintenance; Java; Benchmark (surveying); Security token; Software; Precision and recall; Maintainability; Software system; Programming language; Artificial intelligence; Operating system; Biology; Software engineering; Gene","score_opus":0.01879607424847668,"score_gpt":0.287004374445986,"score_spread":0.2682083001975093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170660163","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16663912,0.00051683385,0.8113228,0.00018833367,0.000056510067,0.00018842996,0.0006648098,0.017785903,0.0026372243],"genre_scores_gemma":[0.50831115,0.00025045944,0.48523638,0.0001247098,0.000039514354,0.00013058835,0.0023371037,0.00091177836,0.0026583488],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9964618,0.00067394617,0.0002942939,0.00076667115,0.001595633,0.00020769647],"domain_scores_gemma":[0.98050827,0.007507857,0.0028735374,0.004123153,0.004750194,0.00023692039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022667504,0.0007181102,0.0009874471,0.0063785333,0.00047839872,0.0019354296,0.0014463031,0.0011653384,0.0011225498],"category_scores_gemma":[0.018981863,0.000338582,0.0008483774,0.004255672,0.0005166383,0.0027020976,0.001763766,0.0010795493,0.0010628399],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052893546,0.0003932312,0.032853458,0.0004947172,0.00021271939,0.0007955914,0.0015703378,0.013816272,0.13170391,0.0060391957,0.005344556,0.80624706],"study_design_scores_gemma":[0.00012516927,0.0009328971,0.043400228,0.0001706761,0.0005691949,0.0038946094,0.0006598905,0.5399105,0.36342412,0.017018346,0.029630322,0.00026410716],"about_ca_topic_score_codex":0.002205984,"about_ca_topic_score_gemma":0.0023348606,"teacher_disagreement_score":0.0063785333,"about_ca_system_score_codex":0.0006344534,"about_ca_system_score_gemma":0.0009335833,"threshold_uncertainty_score":0.011987865},"labels":[],"label_agreement":null},{"id":"W2171172898","doi":"10.1145/1985404.1985407","title":"Extracting code clones for refactoring using combinations of clone metrics","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Victoria","keywords":"Code refactoring; clone (Java method); Computer science; Code (set theory); Programming language; Cloning (programming); Computational biology; Biology; Genetics; Software; DNA","score_opus":0.2161381383002295,"score_gpt":0.3510390307316339,"score_spread":0.1349008924314044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171172898","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3539187,0.0018151005,0.6279698,0.00023928031,0.000063096944,0.001257567,0.002608275,0.009796253,0.0023319458],"genre_scores_gemma":[0.32337233,0.00047611544,0.66870284,0.000055590775,0.000032249623,0.0004918393,0.0052097375,0.0007524556,0.0009069105],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99404824,0.0009569331,0.0011650645,0.0012713007,0.0023651705,0.00019330034],"domain_scores_gemma":[0.96639276,0.014972236,0.0058354037,0.0030885274,0.008923226,0.00078779005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029545927,0.0023271656,0.0023569176,0.01774259,0.0006593528,0.0018999705,0.0011937665,0.0012793137,0.000790303],"category_scores_gemma":[0.027294848,0.0007143183,0.0017534031,0.008386363,0.00043704457,0.0026912587,0.0015194772,0.0008158513,0.00053869194],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045173155,0.00043973263,0.13209768,0.0013632113,0.0006364145,0.0009666542,0.0012882989,0.008465675,0.0671481,0.0016117351,0.002421584,0.7831091],"study_design_scores_gemma":[0.00031568424,0.0021995509,0.2770175,0.0006557781,0.0022260654,0.0063534286,0.0017579147,0.49040982,0.18169785,0.012922027,0.023741776,0.0007026092],"about_ca_topic_score_codex":0.0026263113,"about_ca_topic_score_gemma":0.005057018,"teacher_disagreement_score":0.01774259,"about_ca_system_score_codex":0.0006760709,"about_ca_system_score_gemma":0.0014852748,"threshold_uncertainty_score":0.015625536},"labels":[],"label_agreement":null},{"id":"W2171368158","doi":"10.1109/icsm.2015.7332455","title":"A comparative study on the bug-proneness of different types of code clones","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; clone (Java method); Computer science; Software bug; Code (set theory); Software maintenance; Programming language; Type (biology); Software system; Biology; Software; Genetics; Gene","score_opus":0.13052598699706477,"score_gpt":0.34996060519706074,"score_spread":0.21943461819999596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171368158","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99728894,0.00026151363,0.0018494078,0.000014035438,0.000005573278,0.000032014486,0.000120167904,0.00010217433,0.00032626823],"genre_scores_gemma":[0.9938554,0.00015917182,0.0048245084,0.000012481552,0.000008964484,0.000037012276,0.0006725954,0.00004220047,0.00038768843],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99426115,0.0013290866,0.0007651477,0.0011703171,0.0022064091,0.00026778333],"domain_scores_gemma":[0.87171465,0.08577193,0.015244566,0.008791817,0.016186176,0.002290866],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041971807,0.0005199656,0.00056125806,0.0054813586,0.00052356033,0.00089742866,0.0006036461,0.00069592486,0.0006547709],"category_scores_gemma":[0.046254672,0.00026526317,0.00064865354,0.0026569548,0.0006629182,0.0016931841,0.00075952476,0.00049689866,0.0001388163],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013157215,0.00054886646,0.75969946,0.00080830447,0.0006641231,0.0009116597,0.0057416116,0.0071733645,0.053356748,0.00062033866,0.00078643236,0.16837336],"study_design_scores_gemma":[0.000056580335,0.0026201925,0.9517575,0.000081028425,0.0003980998,0.0022188793,0.0020719592,0.021850113,0.016779412,0.0006010142,0.001462888,0.00010220695],"about_ca_topic_score_codex":0.0013813167,"about_ca_topic_score_gemma":0.0026176623,"teacher_disagreement_score":0.0054813586,"about_ca_system_score_codex":0.00050221843,"about_ca_system_score_gemma":0.00039365157,"threshold_uncertainty_score":0.022197068},"labels":[],"label_agreement":null},{"id":"W2171606859","doi":"10.1109/wcre.2009.41","title":"Who are Source Code Contributors and How do they Change?","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"License; Source code; Computer science; Open source software; Open source; Code (set theory); Empirical research; Software; Computer security; World Wide Web; Code review; Internet privacy; Business; Software quality; Software development; Programming language; Operating system","score_opus":0.02248833389862509,"score_gpt":0.25235615110741333,"score_spread":0.22986781720878824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171606859","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9808065,0.0014127855,0.00450497,0.0035925861,0.00016077331,0.00009253705,0.0003001442,0.00016080412,0.008968759],"genre_scores_gemma":[0.9926801,0.0008660913,0.001874328,0.00022856853,0.00010307993,0.000050363786,0.00024052025,0.00006360692,0.0038934546],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9895243,0.003451844,0.00085608783,0.0017494237,0.0036109267,0.0008075463],"domain_scores_gemma":[0.8752194,0.055286158,0.03815016,0.008300836,0.017486466,0.0055569857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008675928,0.00023642778,0.00050839025,0.0044573657,0.0018157648,0.0042594196,0.000968016,0.0012567119,0.0030621733],"category_scores_gemma":[0.13042323,0.00046396223,0.0002403606,0.0032601152,0.0015659499,0.0078043044,0.0013878423,0.0010510035,0.0016762152],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000081652535,0.00008703728,0.8565507,0.00009423135,0.00006117137,0.00042569666,0.015899,0.00015023013,0.00079161226,0.0010697528,0.0023465992,0.1224424],"study_design_scores_gemma":[0.000011946397,0.000097474236,0.94743615,0.0001305965,0.000060841183,0.0012850013,0.028158784,0.0014095418,0.00084587996,0.0024416875,0.018063556,0.000058557176],"about_ca_topic_score_codex":0.004693908,"about_ca_topic_score_gemma":0.008050491,"teacher_disagreement_score":0.008675928,"about_ca_system_score_codex":0.0011334352,"about_ca_system_score_gemma":0.0012684078,"threshold_uncertainty_score":0.04588324},"labels":[],"label_agreement":null},{"id":"W2171618950","doi":"10.1109/csmr.2005.49","title":"Software Clustering Based on Dynamic Dependencies","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Cluster analysis; Data mining; Software system; Software; Software construction; Software sizing; Software engineering; Artificial intelligence; Programming language","score_opus":0.01310430193199962,"score_gpt":0.2577414525729553,"score_spread":0.24463715064095565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171618950","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5124422,0.0005448541,0.4698864,0.000352272,0.00008672753,0.00039250878,0.0010980265,0.00289661,0.012300481],"genre_scores_gemma":[0.8290916,0.00021968537,0.16663364,0.00004719819,0.000023487202,0.00017935551,0.0016625818,0.00035140407,0.0017909494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954721,0.001060072,0.00027854607,0.00079990586,0.002007673,0.00038182724],"domain_scores_gemma":[0.97737586,0.008913194,0.0020577353,0.004896187,0.0061416985,0.00061528635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029258165,0.0008086174,0.0007832551,0.0069026155,0.0019947696,0.0018441429,0.0012990637,0.0008292371,0.0016393298],"category_scores_gemma":[0.02471644,0.00044127333,0.0007401638,0.0054748333,0.0013860957,0.0032017815,0.0022136497,0.0009754722,0.000476886],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011436092,0.0003901148,0.0646278,0.0006239905,0.00041203355,0.00062674144,0.0019583746,0.38326338,0.033868536,0.051860645,0.007841907,0.45338288],"study_design_scores_gemma":[0.00005104012,0.00022180905,0.023481233,0.000061004586,0.0001388799,0.0005157633,0.00055616844,0.90473336,0.022504723,0.04159157,0.0060251704,0.00011933553],"about_ca_topic_score_codex":0.007815504,"about_ca_topic_score_gemma":0.0095260255,"teacher_disagreement_score":0.007815504,"about_ca_system_score_codex":0.0020815658,"about_ca_system_score_gemma":0.0016994139,"threshold_uncertainty_score":0.015540063},"labels":[],"label_agreement":null},{"id":"W2171683500","doi":"10.5555/2663297.2663305","title":"Generating precise dependencies for large software","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada); University of Waterloo","funders":"","keywords":"Code refactoring; Computer science; Software; Software development; Software engineering; Source lines of code; Software system; Scalability; Code (set theory); Software construction; Source code; Backporting; Programming language; Data mining; Database; Set (abstract data type)","score_opus":0.020952338323343568,"score_gpt":0.26432046299521395,"score_spread":0.24336812467187038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171683500","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17557745,0.00045187204,0.80435663,0.00040940515,0.000072603885,0.00018872558,0.003106914,0.01183854,0.0039978423],"genre_scores_gemma":[0.49298173,0.0002955864,0.49525002,0.000114695096,0.000043431417,0.00021251544,0.0072888033,0.0018879024,0.0019253591],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978224,0.0004308876,0.00017358272,0.0004255893,0.0010332025,0.00011436946],"domain_scores_gemma":[0.98444843,0.008033163,0.0017632822,0.0034109694,0.0021664111,0.0001777996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001664307,0.000732048,0.0004389861,0.0022624978,0.00084353244,0.000844745,0.0010054848,0.0007019299,0.002348784],"category_scores_gemma":[0.016486714,0.0007086476,0.0005920423,0.002226318,0.0006314676,0.0019539376,0.0019723985,0.0014068426,0.0007263238],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003503191,0.00031348754,0.06228956,0.00089753984,0.00015654991,0.0019041916,0.0014720889,0.18994422,0.061959617,0.030245181,0.030234324,0.62023294],"study_design_scores_gemma":[0.000052686424,0.00013171775,0.021554945,0.00014070857,0.00008264395,0.0008712137,0.0002891138,0.83983535,0.060018502,0.05198575,0.024956884,0.000080509],"about_ca_topic_score_codex":0.0026584652,"about_ca_topic_score_gemma":0.004560874,"teacher_disagreement_score":0.0026584652,"about_ca_system_score_codex":0.00068292156,"about_ca_system_score_gemma":0.0014513218,"threshold_uncertainty_score":0.008801818},"labels":[],"label_agreement":null},{"id":"W2171733741","doi":"10.5555/776816.776866","title":"Hipikat: recommending pertinent software development artifacts","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":320,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Eclipse; Computer science; Task (project management); Context (archaeology); Open source; Software project management; Open source software; Software; Open-source software development; Software engineering; Software development; World Wide Web; Data science; Software construction; Engineering; Systems engineering; Programming language","score_opus":0.03155944281280614,"score_gpt":0.2593580953075713,"score_spread":0.22779865249476516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171733741","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23994406,0.0017317198,0.6793337,0.0031307614,0.00033200718,0.0030382124,0.0030121764,0.0517961,0.017681329],"genre_scores_gemma":[0.1535511,0.00059697876,0.83439404,0.00020415208,0.000069460446,0.00047553916,0.002877432,0.0005770252,0.0072543304],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976762,0.0007381993,0.00024688125,0.00047464034,0.00075793115,0.00010614515],"domain_scores_gemma":[0.9858267,0.0078108073,0.0014614806,0.0023999442,0.0017268981,0.0007741594],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044701793,0.0012461921,0.00067985035,0.0046587107,0.0013009817,0.0028525086,0.0022621546,0.0019354755,0.0058087576],"category_scores_gemma":[0.03196399,0.00078877906,0.0005669816,0.002496171,0.00056312006,0.0038607242,0.0024328597,0.0010524933,0.0027303007],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004668004,0.000636174,0.026231732,0.002180508,0.00016298739,0.0011803582,0.009137508,0.0054884604,0.01790811,0.0031890406,0.029624619,0.9037939],"study_design_scores_gemma":[0.00087816524,0.0030812568,0.100789435,0.0023285246,0.0013249478,0.0070748264,0.018784106,0.2181507,0.07742149,0.030058987,0.5390202,0.0010873843],"about_ca_topic_score_codex":0.0047660894,"about_ca_topic_score_gemma":0.016531741,"teacher_disagreement_score":0.0058087576,"about_ca_system_score_codex":0.0006898393,"about_ca_system_score_gemma":0.0028368968,"threshold_uncertainty_score":0.023640871},"labels":[],"label_agreement":null},{"id":"W2171753775","doi":"10.1109/caia.1993.366651","title":"Compositional software reuse with case-based reasoning","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Reuse; Software engineering; Flexibility (engineering); Software; Code reuse; Process (computing); Matching (statistics); Source code; Programming language; Code (set theory); Adaptability; Data mining; Artificial intelligence; Engineering","score_opus":0.019678112325361718,"score_gpt":0.2318524139736261,"score_spread":0.2121743016482644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171753775","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026365197,0.00016841138,0.9893016,0.00037170996,0.000032337975,0.00021387852,0.00002575392,0.00068038836,0.0065693734],"genre_scores_gemma":[0.07324846,0.0003419095,0.9232078,0.00013420811,0.000042747775,0.0003218317,0.00016947625,0.00010011313,0.0024334325],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9893226,0.00517639,0.00083251717,0.00078143,0.0033803442,0.000506756],"domain_scores_gemma":[0.9920442,0.0048286947,0.00047970584,0.001879314,0.0006266504,0.00014138702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00882651,0.001464416,0.0009491696,0.004452294,0.0015764425,0.0065104845,0.0037585301,0.0028615585,0.0069243503],"category_scores_gemma":[0.020800736,0.0016522518,0.0032795102,0.0021333385,0.0058963406,0.0075401007,0.0060495604,0.0026501615,0.0014027455],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008936927,0.00028229158,0.00091495545,0.0004101205,0.00013503047,0.0011142532,0.0020297684,0.10737886,0.0027935146,0.7466222,0.003892744,0.13433692],"study_design_scores_gemma":[0.00012432062,0.000046960973,0.00015962018,0.00018317584,0.00009751869,0.00054077245,0.00027695257,0.2621938,0.004016291,0.7017388,0.030553231,0.00006860433],"about_ca_topic_score_codex":0.0040297937,"about_ca_topic_score_gemma":0.004316673,"teacher_disagreement_score":0.00882651,"about_ca_system_score_codex":0.002093414,"about_ca_system_score_gemma":0.0023240834,"threshold_uncertainty_score":0.046679616},"labels":[],"label_agreement":null},{"id":"W2171758753","doi":"10.1109/caia.1995.378780","title":"Using paths to detect redundancy in rule bases","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Redundancy (engineering); Computer science; Data mining; Inefficiency; Rule-based system; Theoretical computer science; Algorithm; Artificial intelligence","score_opus":0.05865827952763222,"score_gpt":0.2942813950271803,"score_spread":0.2356231154995481,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171758753","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051593322,0.00018933746,0.9423633,0.00020354714,0.000024888684,0.0002354779,0.00022477089,0.00432543,0.0008399303],"genre_scores_gemma":[0.15235372,0.00012872863,0.84567446,0.000061820836,0.000011230926,0.00014550283,0.00056653837,0.00024819674,0.0008098302],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9947876,0.0015137359,0.00048778814,0.0008003463,0.002161363,0.00024920574],"domain_scores_gemma":[0.965789,0.022773085,0.003195526,0.0037349407,0.0042147953,0.00029247993],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004048064,0.0010907389,0.000821423,0.0037173864,0.00071905105,0.002027916,0.0018632457,0.0012273776,0.0021044803],"category_scores_gemma":[0.03175215,0.0009048984,0.0008693592,0.0018296205,0.0014574738,0.0035587347,0.002169506,0.0013268176,0.0007199171],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080244255,0.00021042206,0.018510127,0.0006048547,0.00020164487,0.0007547432,0.0016172046,0.07840637,0.034261182,0.03963968,0.0033778762,0.82161343],"study_design_scores_gemma":[0.00010181408,0.00033337186,0.0035252818,0.0001857297,0.00015719785,0.0010592672,0.00031728714,0.7759117,0.11145356,0.09357308,0.013274105,0.000107529435],"about_ca_topic_score_codex":0.0021702049,"about_ca_topic_score_gemma":0.0021463016,"teacher_disagreement_score":0.004048064,"about_ca_system_score_codex":0.0009622222,"about_ca_system_score_gemma":0.002098939,"threshold_uncertainty_score":0.021408439},"labels":[],"label_agreement":null},{"id":"W2171868993","doi":"10.1109/wcre.2010.11","title":"Studying the Impact of Clones on Software Defects","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Cloning (programming); Commit; Computer science; Software; Code (set theory); Software maintenance; Software engineering; Software system; Software evolution; Programming language; Software construction; Biology; Genetics; Database; Gene","score_opus":0.023275071991492504,"score_gpt":0.30559533710745274,"score_spread":0.28232026511596026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171868993","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.984287,0.00024856915,0.0147440145,0.00006059367,0.0000037874547,0.000013787466,0.000057276604,0.000075810814,0.00050920085],"genre_scores_gemma":[0.9966468,0.0001235167,0.0029167116,0.000009996731,0.000005373365,0.000009436415,0.00007820039,0.000012099034,0.00019777936],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9982256,0.00061341084,0.000076269185,0.00026011883,0.0006724288,0.00015206989],"domain_scores_gemma":[0.86149484,0.11775096,0.011951035,0.0032616928,0.004535206,0.0010062652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003006413,0.00050708355,0.00044704904,0.0021970575,0.00026446872,0.0006244658,0.00044027777,0.0006694519,0.0009214494],"category_scores_gemma":[0.03635502,0.00023193705,0.00071128696,0.001100634,0.00075297494,0.0014246887,0.0005867486,0.0007183555,0.0001158537],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039794334,0.00025941967,0.7448582,0.00025874193,0.00052305375,0.00078148645,0.0009862761,0.13489558,0.01742926,0.0034282836,0.00028662581,0.09589517],"study_design_scores_gemma":[0.00002775176,0.0018066028,0.541272,0.000075759155,0.0003711774,0.0011542316,0.00065010163,0.42900187,0.016594525,0.008194409,0.0007697916,0.000081858714],"about_ca_topic_score_codex":0.0027269593,"about_ca_topic_score_gemma":0.0026767752,"teacher_disagreement_score":0.003006413,"about_ca_system_score_codex":0.00062528177,"about_ca_system_score_gemma":0.0005519828,"threshold_uncertainty_score":0.015899599},"labels":[],"label_agreement":null},{"id":"W2171932879","doi":"10.1109/case.1992.200160","title":"A framework for capturing design rationale using granularity hierarchies","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Guangdong Academy of Sciences","keywords":"Granularity; Schematic; Computer science; Usability; Notation; Component (thermodynamics); Software; Software engineering; Theoretical computer science; Programming language; Data mining; Human–computer interaction; Mathematics; Engineering","score_opus":0.0966438150107016,"score_gpt":0.316822960167619,"score_spread":0.22017914515691742,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171932879","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00065689924,0.0003265242,0.9928888,0.0005406514,0.00005297325,0.00019970688,0.00019560807,0.0010505507,0.0040880996],"genre_scores_gemma":[0.013075814,0.0003857174,0.9840059,0.00015170754,0.000035919686,0.0003492811,0.00038706718,0.00013401035,0.0014745236],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9867286,0.0057753026,0.0021230495,0.0014419397,0.003345855,0.0005852158],"domain_scores_gemma":[0.98525697,0.006845468,0.0015241972,0.0035146398,0.0022785773,0.0005802049],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015900025,0.0024109238,0.0011624071,0.008244055,0.0030367074,0.010594194,0.003675574,0.0038006087,0.0063916626],"category_scores_gemma":[0.026551444,0.002499109,0.0033198916,0.006504331,0.00630928,0.017930767,0.004589468,0.006121429,0.003257225],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026881326,0.000040279825,0.00036995593,0.00036118022,0.00003512733,0.00021540026,0.0026748553,0.0050030933,0.0016025188,0.91364145,0.004751137,0.0712781],"study_design_scores_gemma":[0.00004756033,0.000057566114,0.00027596264,0.00078342116,0.00007425266,0.0005300258,0.000784637,0.03737221,0.0019172508,0.78775316,0.17029615,0.0001077982],"about_ca_topic_score_codex":0.01342014,"about_ca_topic_score_gemma":0.013897759,"teacher_disagreement_score":0.015900025,"about_ca_system_score_codex":0.0040362617,"about_ca_system_score_gemma":0.0072685964,"threshold_uncertainty_score":0.084088385},"labels":[],"label_agreement":null},{"id":"W2172101821","doi":"10.1109/ccece.2002.1013033","title":"Release date prediction for telecommunication software using Bayesian Belief Networks","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; University of Calgary","keywords":"Computer science; Bayesian network; Data mining; Context (archaeology); Process (computing); Schedule; Set (abstract data type); Software quality; Code (set theory); Source code; Software; Variable (mathematics); Quality (philosophy); Software metric; Machine learning; Software development","score_opus":0.024877061414091767,"score_gpt":0.27405390170189525,"score_spread":0.24917684028780349,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2172101821","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40186355,0.00071979896,0.59338623,0.0006704026,0.00004594875,0.00010395625,0.0003992847,0.00094547373,0.0018653106],"genre_scores_gemma":[0.9314102,0.0003406576,0.06653054,0.00004236406,0.000027441152,0.0000791372,0.00049601874,0.000045808218,0.001027809],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99861634,0.0006179306,0.00009781717,0.00019619065,0.00035586164,0.00011587321],"domain_scores_gemma":[0.98513216,0.012108054,0.0011719553,0.00027352304,0.001127791,0.00018659643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051657143,0.001074361,0.0009650104,0.0029285455,0.00055196125,0.0015849151,0.0011199729,0.0012936488,0.001072081],"category_scores_gemma":[0.019517532,0.0010119409,0.0008254302,0.00169526,0.0005681111,0.002036616,0.00066858926,0.0013241059,0.00029801158],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025526382,0.00008215822,0.005706809,0.00003288945,0.00004674912,0.000034871282,0.000070818256,0.9649617,0.00034396542,0.0012993927,0.00033861015,0.026826821],"study_design_scores_gemma":[0.0000057282828,0.000008950775,0.00044443668,0.0000028646343,0.000004602032,0.0000030047147,0.000003658989,0.99859184,0.00011180817,0.00078620855,0.000032725184,0.0000041962244],"about_ca_topic_score_codex":0.050570156,"about_ca_topic_score_gemma":0.034991246,"teacher_disagreement_score":0.050570156,"about_ca_system_score_codex":0.0027595889,"about_ca_system_score_gemma":0.0011216967,"threshold_uncertainty_score":0.100551605},"labels":[],"label_agreement":null},{"id":"W2172298781","doi":"10.1109/sess.1995.525960","title":"Ontario Hydro/AECL standards for software engineering deficiencies in existing standards that created their need","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Hydro One (Canada)","funders":"","keywords":"Engineering; Software; Systems engineering; Engineering management; Computer science","score_opus":0.04962133625241198,"score_gpt":0.2660631625758955,"score_spread":0.21644182632348352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2172298781","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024222832,0.0060596094,0.2603661,0.114023,0.0027012634,0.0028348442,0.0042051594,0.010115128,0.5754722],"genre_scores_gemma":[0.18738452,0.0075018234,0.5311515,0.011560873,0.00074330723,0.0023914673,0.007860179,0.0031484673,0.24825795],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9255395,0.004764159,0.0048616114,0.0020038327,0.059154876,0.0036759947],"domain_scores_gemma":[0.74881643,0.014695212,0.0074757696,0.019035053,0.204103,0.0058744834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026960181,0.00068194175,0.0006801037,0.007607579,0.007435203,0.00928225,0.0068749785,0.0038342532,0.010509561],"category_scores_gemma":[0.08199974,0.0010996881,0.0009408439,0.010265355,0.005467996,0.0072005135,0.004019793,0.0034393014,0.005453714],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000118737764,0.00018856519,0.0076738866,0.0010638282,0.000029537623,0.0003316126,0.0075514875,0.0031963126,0.007506863,0.36930948,0.35644406,0.24658562],"study_design_scores_gemma":[0.000034815923,0.00003768069,0.005134,0.0006360291,0.000023101413,0.00016579236,0.0015468247,0.0034136858,0.0024607796,0.01635054,0.9701104,0.000086411026],"about_ca_topic_score_codex":0.7984989,"about_ca_topic_score_gemma":0.8494252,"teacher_disagreement_score":0.20150107,"about_ca_system_score_codex":0.066827446,"about_ca_system_score_gemma":0.1694489,"threshold_uncertainty_score":0.4848693},"labels":[],"label_agreement":null},{"id":"W2172820611","doi":"10.1007/978-3-319-24318-4_6","title":"SATGraf: Visualizing the Evolution of SAT Formula Structure in Solvers","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Boolean satisfiability problem; Maximum satisfiability problem; Heuristic; Branching (polymer chemistry); Theoretical computer science; Solver; Satisfiability; Boolean data type; Heuristics; True quantified Boolean formula; Algorithm; Boolean function; Artificial intelligence; Programming language","score_opus":0.023222380151059263,"score_gpt":0.2810982460035222,"score_spread":0.2578758658524629,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2172820611","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05638147,0.002636024,0.7135773,0.0024849027,0.00073753315,0.00019102945,0.03032807,0.15174565,0.041918088],"genre_scores_gemma":[0.27789122,0.0015134637,0.6538847,0.0006012027,0.000111225054,0.0003491972,0.029133527,0.020027645,0.016487792],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995763,0.00012778823,0.00002516458,0.00007797456,0.00014587057,0.00004699215],"domain_scores_gemma":[0.99786264,0.0014080587,0.000102236416,0.00021200917,0.0003036202,0.000111583184],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00068879215,0.00111919,0.0005440792,0.0021197523,0.00060452905,0.002912544,0.0016656964,0.0012008059,0.045775983],"category_scores_gemma":[0.004884196,0.0006064984,0.00088151585,0.0019285373,0.0004364384,0.0028955608,0.0013704149,0.0024235882,0.005063252],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093264243,0.00027663365,0.008320312,0.0019299962,0.00018817674,0.00064774277,0.003283006,0.10068183,0.02589371,0.114853375,0.36915538,0.37383714],"study_design_scores_gemma":[0.00020870079,0.00011528333,0.0034060585,0.00036114297,0.000072926756,0.00032989832,0.00068089174,0.7156905,0.022843432,0.10798163,0.14820246,0.0001071165],"about_ca_topic_score_codex":0.0071167136,"about_ca_topic_score_gemma":0.01272632,"teacher_disagreement_score":0.045775983,"about_ca_system_score_codex":0.00080302,"about_ca_system_score_gemma":0.001186623,"threshold_uncertainty_score":0.1531359},"labels":[],"label_agreement":null},{"id":"W2173092635","doi":"10.1109/models.2015.7338235","title":"State machine antipatterns for UML-RT","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software engineering; Software development; Unified Modeling Language; Software construction; Software; Software development process; Software system; Set (abstract data type); Software quality; Programming language","score_opus":0.04992134072962341,"score_gpt":0.3129910642982242,"score_spread":0.2630697235686008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2173092635","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030177997,0.00015164155,0.98013765,0.00025448372,0.00010593579,0.00014536822,0.00026524757,0.008742529,0.0071793515],"genre_scores_gemma":[0.107027516,0.00044520473,0.8805163,0.0004716631,0.00009922338,0.000756817,0.0010911549,0.0025862672,0.007005978],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.994785,0.001594705,0.00073166145,0.00057756057,0.002086698,0.00022434973],"domain_scores_gemma":[0.99322635,0.002830238,0.000997442,0.0018850206,0.0009796972,0.00008136613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026611595,0.0009159562,0.00056260504,0.0013841755,0.0005821238,0.001924156,0.0011887703,0.0016683756,0.0075196936],"category_scores_gemma":[0.009853833,0.0006657521,0.0010886623,0.00083474105,0.0018551445,0.002681075,0.0015927909,0.002641556,0.003223362],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002590242,0.00015206235,0.0021700226,0.00071558147,0.000051070798,0.0008210345,0.0019755575,0.01990995,0.027310627,0.5732634,0.020736603,0.35263512],"study_design_scores_gemma":[0.000086707136,0.00023758515,0.0010514457,0.00040552075,0.00006992309,0.0013154137,0.00016395579,0.16818443,0.04546835,0.3199305,0.46295705,0.00012917453],"about_ca_topic_score_codex":0.0014913155,"about_ca_topic_score_gemma":0.0020732097,"teacher_disagreement_score":0.0075196936,"about_ca_system_score_codex":0.0007657821,"about_ca_system_score_gemma":0.0012228802,"threshold_uncertainty_score":0.025155902},"labels":[],"label_agreement":null},{"id":"W2175159659","doi":"","title":"Recommending software experts using code similarity and social heuristics","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Recommender system; Heuristics; Software; Code (set theory); Code review; Software engineering; Similarity (geometry); World Wide Web; Software development; Work (physics); Software quality; Data science; Engineering; Artificial intelligence; Set (abstract data type)","score_opus":0.05266931925355162,"score_gpt":0.3070951386637148,"score_spread":0.2544258194101632,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2175159659","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7851312,0.001099307,0.1982876,0.0007049403,0.00010636101,0.0006694439,0.0009810367,0.0022357462,0.010784377],"genre_scores_gemma":[0.88532525,0.00018074433,0.1103236,0.0001106783,0.00005905827,0.00012396103,0.0009621112,0.00005736928,0.002857292],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984459,0.0004989368,0.00009581413,0.0004271888,0.00042367043,0.000108437125],"domain_scores_gemma":[0.9928429,0.0046411282,0.00057006493,0.00042853496,0.0012034789,0.00031383758],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024904595,0.0010126479,0.00084215635,0.0054510552,0.0007390851,0.0013190277,0.001181408,0.0012885997,0.0014922973],"category_scores_gemma":[0.01106404,0.0004027391,0.0005455241,0.0018284693,0.00031908706,0.0015024722,0.0007098774,0.00061928184,0.00081229076],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00163202,0.0017795603,0.18447016,0.0006374595,0.00072688394,0.00050067174,0.0012598116,0.11072115,0.010064704,0.004304709,0.016717795,0.66718507],"study_design_scores_gemma":[0.00013917655,0.00037453198,0.021897094,0.000031834184,0.00013056259,0.00020254684,0.00034924873,0.9688665,0.0027304802,0.0023553064,0.0028650248,0.00005769607],"about_ca_topic_score_codex":0.017593995,"about_ca_topic_score_gemma":0.034314286,"teacher_disagreement_score":0.017593995,"about_ca_system_score_codex":0.00092818675,"about_ca_system_score_gemma":0.0010040705,"threshold_uncertainty_score":0.034983158},"labels":[],"label_agreement":null},{"id":"W2177378206","doi":"10.1007/978-3-319-24285-9_1","title":"A Suite of Rules for Developing and Evaluating Software Quality Models","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software quality; Quality (philosophy); Software quality control; Computer science; Software quality analyst; Verification and validation; Ambiguity; Software; Suite; Software metric; Set (abstract data type); Software engineering; Data mining; Software development; Reliability engineering; Engineering; Operations management; Programming language; Geography","score_opus":0.10624179440377576,"score_gpt":0.352118014418156,"score_spread":0.24587622001438023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2177378206","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0039829756,0.00021960361,0.9880874,0.00021650785,0.00004536053,0.00067823316,0.00082669145,0.0029194076,0.0030238742],"genre_scores_gemma":[0.033370312,0.0002568592,0.9628641,0.00012286687,0.000036253627,0.0006415748,0.0015116839,0.00032779967,0.0008684742],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9651217,0.006748769,0.0070511014,0.0025943362,0.017586105,0.00089803466],"domain_scores_gemma":[0.94641167,0.029023161,0.0041002664,0.0063427724,0.013449988,0.00067213294],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018984571,0.0023654262,0.002953635,0.011082587,0.0017136845,0.008086227,0.0041179783,0.002866086,0.0037306512],"category_scores_gemma":[0.08638655,0.0015336244,0.004758588,0.0047664694,0.0013966232,0.006675793,0.003264188,0.0035381955,0.0025376389],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018600794,0.0006491132,0.011145091,0.0011458173,0.00051921577,0.0006223196,0.00054361374,0.10705782,0.0053610816,0.10169398,0.019053312,0.7520226],"study_design_scores_gemma":[0.00008055136,0.0002211067,0.0025768166,0.0009588728,0.0005582475,0.0004834972,0.00024616226,0.7774897,0.014877082,0.18338719,0.018970452,0.0001503856],"about_ca_topic_score_codex":0.0067649274,"about_ca_topic_score_gemma":0.010056294,"teacher_disagreement_score":0.018984571,"about_ca_system_score_codex":0.002150015,"about_ca_system_score_gemma":0.00402188,"threshold_uncertainty_score":0.10040122},"labels":[],"label_agreement":null},{"id":"W2182913975","doi":"","title":"LEVENSHTEIN EDIT DISTANCE-BASED TYPE III CLONE DETECTION USING METRIC TREES","year":2011,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Levenshtein distance; Metric (unit); clone (Java method); Computer science; Edit distance; Code (set theory); Mathematics; Artificial intelligence; Algorithm; Biology; Engineering; Programming language; Genetics; DNA","score_opus":0.02734475168536953,"score_gpt":0.24424384653310616,"score_spread":0.21689909484773662,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2182913975","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06663955,0.00039168698,0.9281636,0.00007852438,0.00004844913,0.000078914614,0.00018405048,0.0032099483,0.0012053307],"genre_scores_gemma":[0.37302575,0.00016701272,0.6234148,0.00006551824,0.000039162143,0.0000817235,0.0007798181,0.00031696493,0.002109207],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9965461,0.0005790434,0.00029649734,0.0005747732,0.0018343983,0.00016923241],"domain_scores_gemma":[0.9922172,0.0024517474,0.0010113542,0.0013877318,0.0026764183,0.00025550727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014141276,0.00070069276,0.0008280578,0.0033921266,0.000569505,0.0014519494,0.0015525299,0.0010128109,0.0013341652],"category_scores_gemma":[0.009022518,0.00026354167,0.00073231914,0.0026724301,0.00066810224,0.001778699,0.0013458648,0.00082653324,0.0009154634],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037561203,0.00017608367,0.013508835,0.0004105068,0.000215632,0.00034821738,0.00055814686,0.025091596,0.111466184,0.009611038,0.003169505,0.8350686],"study_design_scores_gemma":[0.00003776482,0.00058649015,0.011424514,0.000047281315,0.00010182597,0.0023006538,0.00019435653,0.74539983,0.21409419,0.014400919,0.011298481,0.00011371842],"about_ca_topic_score_codex":0.0021629825,"about_ca_topic_score_gemma":0.0021760748,"teacher_disagreement_score":0.0033921266,"about_ca_system_score_codex":0.000733323,"about_ca_system_score_gemma":0.0008247603,"threshold_uncertainty_score":0.0074787736},"labels":[],"label_agreement":null},{"id":"W2183869536","doi":"10.22215/etd/2012-09641","title":"An approach to design pattern and anti-pattern detection in MOF-based modeling languages","year":2012,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Canadian Heritage; Library and Archives Canada","funders":"","keywords":"Computer science; Humanities; Art","score_opus":0.03293593852257042,"score_gpt":0.2967618130366783,"score_spread":0.2638258745141079,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2183869536","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011778901,0.0000467722,0.9963456,0.00018116312,0.000023241415,0.00008346607,0.00004618415,0.0012885579,0.00080711104],"genre_scores_gemma":[0.01646472,0.00008927134,0.9808076,0.00013908505,0.000018587276,0.00012823899,0.00016780235,0.00038346054,0.0018012106],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954189,0.0013468664,0.00058500376,0.0005917817,0.0017562562,0.0003011933],"domain_scores_gemma":[0.9937499,0.0026451072,0.00045974232,0.0017776015,0.0011780363,0.00018955566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054397653,0.0009796554,0.0007543338,0.0020837064,0.0011395654,0.0042999573,0.003533508,0.0024828115,0.004566985],"category_scores_gemma":[0.009360717,0.0016479149,0.0030418206,0.001354946,0.0023081615,0.0055123344,0.0032742275,0.0039598094,0.0014651861],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002130619,0.000388242,0.0026464825,0.0006475268,0.00023357086,0.0005931965,0.0031510273,0.033863742,0.020896275,0.5382012,0.009473447,0.38969213],"study_design_scores_gemma":[0.00010922043,0.00023098031,0.00048337283,0.00035592495,0.00024873976,0.00093879586,0.00069663627,0.5815122,0.034285914,0.23768674,0.14329624,0.00015517905],"about_ca_topic_score_codex":0.004864596,"about_ca_topic_score_gemma":0.010125349,"teacher_disagreement_score":0.0054397653,"about_ca_system_score_codex":0.0014226286,"about_ca_system_score_gemma":0.002649475,"threshold_uncertainty_score":0.02876854},"labels":[],"label_agreement":null},{"id":"W2185225911","doi":"10.1007/978-3-642-15405-8_6","title":"Programming Expertise during Incremental Software Development: An Empirical Study","year":2010,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Algoma University","funders":"Algoma University","keywords":"Debugging; Computer science; Taxonomy (biology); Domain (mathematical analysis); Bloom's taxonomy; Software engineering; Software development; Subject-matter expert; Domain knowledge; Software; Process (computing); Data science; Knowledge management; Artificial intelligence; Cognition; Programming language; Expert system; Psychology","score_opus":0.11976024379682554,"score_gpt":0.40064254989710085,"score_spread":0.2808823061002753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2185225911","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9978218,0.00010134751,0.00035013558,0.00003625386,0.0000022568772,0.000022712804,0.000012722819,0.000007224267,0.0016455025],"genre_scores_gemma":[0.9985476,0.00010812287,0.00047865,0.000023936906,0.0000058745222,0.00003206578,0.000050296763,0.000007885968,0.00074547954],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99632156,0.0019067388,0.00021298886,0.00037773076,0.00083701796,0.00034400224],"domain_scores_gemma":[0.76371753,0.20195569,0.014151074,0.006836445,0.0072110975,0.006128139],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.006065313,0.00031714304,0.00035798125,0.0013099026,0.0011774276,0.0011823418,0.0014411507,0.0010972891,0.0036223289],"category_scores_gemma":[0.093989104,0.000501084,0.00024933807,0.00090785587,0.0011844407,0.0025270574,0.0016400566,0.0020720402,0.0005849673],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00241902,0.018161092,0.75296915,0.00045690002,0.000081393315,0.0017340656,0.08056077,0.002186032,0.0039133416,0.0020142773,0.0013664843,0.13413745],"study_design_scores_gemma":[0.00028205238,0.0068043596,0.92339057,0.00031097568,0.00019670889,0.0029721726,0.042673856,0.010563335,0.004201356,0.003884191,0.0045988164,0.000121622325],"about_ca_topic_score_codex":0.003622274,"about_ca_topic_score_gemma":0.004577478,"teacher_disagreement_score":0.9939347,"about_ca_system_score_codex":0.00078806945,"about_ca_system_score_gemma":0.0014266649,"threshold_uncertainty_score":0.032076776},"labels":[],"label_agreement":null},{"id":"W2185471824","doi":"","title":"An exploratory study of the impact of software changeability","year":2009,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Eclipse; Software; Computer science; Software quality; Quality (philosophy); Exploratory research; Test (biology); Software engineering; Software development; Biology; Programming language","score_opus":0.018415653487793407,"score_gpt":0.27737877744819944,"score_spread":0.25896312396040605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2185471824","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9990815,0.000019995236,0.0003480619,0.0000204216,0.0000011557581,0.000020856118,0.00006069937,0.000013022259,0.00043418127],"genre_scores_gemma":[0.9991968,0.000011222784,0.00055683265,0.000010639146,0.0000024326496,0.000015238071,0.000091233414,0.0000066035645,0.00010901605],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9960841,0.0021475162,0.00020483193,0.00043067298,0.0009064272,0.00022640634],"domain_scores_gemma":[0.885324,0.09151136,0.012648114,0.004139971,0.0047670086,0.0016095538],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034536102,0.00027057485,0.00020447043,0.0014602322,0.00045661005,0.0008368846,0.00057884905,0.0005102922,0.001449264],"category_scores_gemma":[0.029081777,0.00023360757,0.00029037974,0.0010590088,0.0005570854,0.0011386913,0.0007533083,0.00075529405,0.00017560169],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008084464,0.0013422242,0.95703596,0.000128505,0.000113401795,0.00084246986,0.0060719093,0.00072832586,0.009639636,0.00034490222,0.0003430053,0.022601226],"study_design_scores_gemma":[0.000036302587,0.0015267687,0.9862874,0.000020477364,0.000045656,0.0005346686,0.002582896,0.004317025,0.0035586208,0.0002339359,0.00083694217,0.000019433503],"about_ca_topic_score_codex":0.0017271143,"about_ca_topic_score_gemma":0.0022212374,"teacher_disagreement_score":0.0034536102,"about_ca_system_score_codex":0.00053295144,"about_ca_system_score_gemma":0.00035604046,"threshold_uncertainty_score":0.018264651},"labels":[],"label_agreement":null},{"id":"W2189543840","doi":"","title":"Using structural relationships to facilitate api learning","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Application programming interface; Reuse; Context (archaeology); World Wide Web; Process (computing); Software engineering; Software; Programming by demonstration; Multimedia; Human–computer interaction; Artificial intelligence; Programming language; Engineering","score_opus":0.20643097360876214,"score_gpt":0.32830590187704267,"score_spread":0.12187492826828053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2189543840","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18780786,0.0007353005,0.75160754,0.0025898279,0.00013219318,0.0009773875,0.00030101763,0.009462366,0.04638655],"genre_scores_gemma":[0.39380386,0.0006308084,0.5919507,0.0005372024,0.000043313732,0.0008644532,0.00082595565,0.0013725834,0.009971132],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9956564,0.0021781577,0.0002458707,0.0006708935,0.0009808815,0.0002677293],"domain_scores_gemma":[0.96346295,0.02679879,0.0025420273,0.003908723,0.0025627576,0.0007247134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044276887,0.0010440124,0.000343194,0.001473515,0.0011545578,0.0036763304,0.0024314218,0.0014192774,0.010955562],"category_scores_gemma":[0.050223086,0.0007230606,0.00076706684,0.0009784847,0.0015026239,0.011516777,0.004843996,0.0023190582,0.0026739698],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030572186,0.0014403118,0.015741155,0.0021296733,0.00005787959,0.0012874502,0.050805047,0.004272106,0.04504409,0.08462176,0.016969718,0.7773251],"study_design_scores_gemma":[0.0003045933,0.0022097381,0.022293165,0.0017947902,0.00038237256,0.004504726,0.02184685,0.091487974,0.07644803,0.15375735,0.6245539,0.00041640992],"about_ca_topic_score_codex":0.0010585836,"about_ca_topic_score_gemma":0.0025513833,"teacher_disagreement_score":0.010955562,"about_ca_system_score_codex":0.0009455601,"about_ca_system_score_gemma":0.0019213788,"threshold_uncertainty_score":0.03665006},"labels":[],"label_agreement":null},{"id":"W2196997841","doi":"","title":"Eight maxims for software inspectors: Research Articles","year":2004,"lang":"en","type":"article","venue":"Software Testing Verification and Reliability","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Cover (algebra); Set (abstract data type); Software; Software inspection; Engineering; Engineering ethics; Computer science; Software engineering; Engineering management; Sociology; Software development; Software quality; Mechanical engineering","score_opus":0.07757743155948835,"score_gpt":0.3266650500084117,"score_spread":0.24908761844892333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2196997841","genre_codex":"review","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017194394,0.48985675,0.009505425,0.37939814,0.02922897,0.0007962886,0.0005927523,0.00026510408,0.07316216],"genre_scores_gemma":[0.16380355,0.63740927,0.032941915,0.055591222,0.040801983,0.0024987538,0.001929757,0.00049348304,0.06453008],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9827717,0.0064336834,0.0025886616,0.00091175875,0.0065327934,0.00076139404],"domain_scores_gemma":[0.8497234,0.10649502,0.009047345,0.0027170433,0.025712345,0.006304829],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023365673,0.0010385851,0.0011484043,0.011748913,0.0040920186,0.01914377,0.001456853,0.0058302903,0.005629617],"category_scores_gemma":[0.09638973,0.0008687281,0.00060178747,0.017830074,0.005342924,0.01652768,0.00512289,0.005475659,0.0023098015],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026520228,0.00033525765,0.0019027079,0.017143019,0.000073345356,0.0003784376,0.02239866,0.00032908388,0.001793244,0.094792105,0.38674968,0.47383928],"study_design_scores_gemma":[0.000049930186,0.00016679673,0.0034445687,0.025068585,0.000079241625,0.00045676876,0.0332283,0.00022141992,0.0006697369,0.037262708,0.899288,0.00006398183],"about_ca_topic_score_codex":0.0006853245,"about_ca_topic_score_gemma":0.0010858141,"teacher_disagreement_score":0.023365673,"about_ca_system_score_codex":0.0061269873,"about_ca_system_score_gemma":0.0074224295,"threshold_uncertainty_score":0.12357098},"labels":[],"label_agreement":null},{"id":"W2199254366","doi":"10.7287/peerj.preprints.1597v1","title":"Multi-token code suggestions using statistical language models","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Naturalness; Security token; Computer science; Surprise; Programmer; Code (set theory); Metric (unit); Simple (philosophy); Programming language; Psychology; Computer security","score_opus":0.1135364606180074,"score_gpt":0.3599472469151834,"score_spread":0.24641078629717603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2199254366","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04116384,0.00015186978,0.9202168,0.00070775184,0.00013928355,0.00009880037,0.0005363196,0.03587763,0.0011076774],"genre_scores_gemma":[0.32075843,0.00009259662,0.67288125,0.00022546746,0.00007382813,0.00012958724,0.0008558543,0.0022297855,0.0027531707],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99757224,0.0012043285,0.00011200056,0.0004894623,0.00052063074,0.00010128014],"domain_scores_gemma":[0.9772379,0.017273841,0.0012600811,0.0017764891,0.0020831747,0.00036846707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034474323,0.0011745951,0.0007214476,0.0011734483,0.00058782444,0.001519067,0.0021725667,0.0010044379,0.0045832093],"category_scores_gemma":[0.02331807,0.00060961314,0.0008971594,0.0006933893,0.0006792027,0.0030845175,0.0011701153,0.0020995021,0.002367064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024870164,0.0005519444,0.015722096,0.0011037878,0.0003074935,0.0013725999,0.0019005489,0.26089752,0.036670636,0.032178864,0.029988939,0.61681855],"study_design_scores_gemma":[0.00003788641,0.00006738034,0.000362945,0.000013007591,0.000018646557,0.00009180424,0.000053650503,0.9794728,0.0065825256,0.010346764,0.0029202886,0.000032276115],"about_ca_topic_score_codex":0.0044920403,"about_ca_topic_score_gemma":0.012730773,"teacher_disagreement_score":0.0045832093,"about_ca_system_score_codex":0.00076886016,"about_ca_system_score_gemma":0.0018655502,"threshold_uncertainty_score":0.018231988},"labels":[],"label_agreement":null},{"id":"W2200054515","doi":"10.1007/s10664-015-9409-1","title":"On the unreliability of bug severity data","year":2015,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Waterloo","funders":"","keywords":"Computer science; Software bug; Reliability (semiconductor); Eclipse; Software; Data mining; Data science; Programming language","score_opus":0.10725834666028859,"score_gpt":0.32480638545586177,"score_spread":0.21754803879557316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2200054515","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6203692,0.005044597,0.33046567,0.020437956,0.0007031248,0.0003131246,0.006355788,0.0018721076,0.014438348],"genre_scores_gemma":[0.9758321,0.0004980472,0.01894347,0.0012158997,0.00035172404,0.00009840465,0.0022613965,0.00021118391,0.0005877369],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.8796456,0.07537611,0.0092413025,0.0148765715,0.019089129,0.0017712824],"domain_scores_gemma":[0.10724069,0.80469483,0.0278662,0.0482057,0.01117944,0.0008131831],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09695513,0.00081707706,0.0016303252,0.009304969,0.0016863894,0.0034945388,0.003583395,0.003928188,0.0029721402],"category_scores_gemma":[0.65337944,0.0014451458,0.0013667003,0.008741851,0.0052410383,0.007333129,0.0031392048,0.0054550176,0.00065433263],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010354852,0.00030238004,0.63106614,0.001294127,0.002069394,0.0017333252,0.0040858597,0.10581302,0.001272804,0.08234027,0.019929111,0.14905816],"study_design_scores_gemma":[0.00036885394,0.00037174617,0.25130937,0.0020633088,0.0009832568,0.0038464963,0.0021487984,0.3349619,0.0046092654,0.3798098,0.019250115,0.0002771238],"about_ca_topic_score_codex":0.006121226,"about_ca_topic_score_gemma":0.0045716045,"teacher_disagreement_score":0.9030449,"about_ca_system_score_codex":0.0023490533,"about_ca_system_score_gemma":0.0019229101,"threshold_uncertainty_score":0.51275384},"labels":[],"label_agreement":null},{"id":"W2202127834","doi":"10.1109/iccd.2015.7357081","title":"Clustering-based revision debug in regression verification","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Debugging; Computer science; Ranking (information retrieval); Cluster analysis; Overhead (engineering); Automation; Rank (graph theory); Data mining; Algorithmic program debugging; Machine learning; Reliability engineering; Programming language; Engineering","score_opus":0.051099181493584105,"score_gpt":0.3141047332282426,"score_spread":0.2630055517346585,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2202127834","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11143319,0.0006390064,0.8792111,0.00024815951,0.0000754185,0.00019789756,0.0001700046,0.0063490984,0.0016760427],"genre_scores_gemma":[0.6715647,0.00015022828,0.3257863,0.00008143029,0.000035276087,0.000100543206,0.0003715782,0.0003331662,0.0015768033],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99604046,0.0015368215,0.0002513783,0.0007632165,0.0010254674,0.0003827544],"domain_scores_gemma":[0.9853498,0.0065429164,0.0020003116,0.0029915571,0.002646946,0.00046853587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005124276,0.0010281105,0.0015374215,0.0037912033,0.0011629515,0.0013999062,0.0023319721,0.0011685453,0.0015338054],"category_scores_gemma":[0.022786597,0.00059125083,0.0008575664,0.0020047466,0.0011773002,0.0020086246,0.0017999627,0.001138404,0.00062838354],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006725324,0.00028606816,0.017263157,0.00019418231,0.00011917046,0.00024398918,0.0005781861,0.49023837,0.008636047,0.013562089,0.002896254,0.46530998],"study_design_scores_gemma":[0.000037958787,0.0001585824,0.0017593537,0.000021876973,0.000031782154,0.00011603045,0.000103016326,0.9817831,0.0056366012,0.008888498,0.0014303459,0.000032946253],"about_ca_topic_score_codex":0.005839731,"about_ca_topic_score_gemma":0.009450653,"teacher_disagreement_score":0.005839731,"about_ca_system_score_codex":0.001270683,"about_ca_system_score_gemma":0.0023200712,"threshold_uncertainty_score":0.027100086},"labels":[],"label_agreement":null},{"id":"W2202473840","doi":"","title":"Monitoring sentiment in open source mailing lists: exploratory study on the apache ecosystem","year":2014,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Sentiment analysis; Software; World Wide Web; Open source; Happiness; Software engineering; Data science; Artificial intelligence; Operating system; Psychology","score_opus":0.02911306773837615,"score_gpt":0.26200637750892514,"score_spread":0.23289330977054898,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2202473840","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9992963,0.000012316375,0.00023946962,0.000034327226,0.0000023456184,0.00002104263,0.00006495862,0.000010398057,0.00031883502],"genre_scores_gemma":[0.9976204,0.000040482588,0.0010943463,0.00006161536,0.000017936633,0.000070467744,0.00039464395,0.000025043826,0.00067497627],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99768996,0.0011894702,0.000119595876,0.00022564267,0.00056108943,0.0002141907],"domain_scores_gemma":[0.97751015,0.014046655,0.0036376496,0.00057172944,0.0031370253,0.0010967866],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034719985,0.00030437004,0.00034391208,0.001513563,0.0012267021,0.0011525325,0.00034971305,0.00054321333,0.00062753307],"category_scores_gemma":[0.015158509,0.00018365197,0.0002596571,0.0011428059,0.00060884666,0.0012944451,0.0008380289,0.00069621543,0.00032291087],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006506576,0.002409881,0.77203983,0.00042236297,0.00009839446,0.0017132845,0.134264,0.00049964373,0.027816502,0.00044279598,0.0036225345,0.056020144],"study_design_scores_gemma":[0.000016037518,0.0005995555,0.93640333,0.00005336426,0.000037422076,0.00033068063,0.052016076,0.0037500786,0.0027413408,0.00020118787,0.0038044795,0.00004643695],"about_ca_topic_score_codex":0.0029175347,"about_ca_topic_score_gemma":0.00535722,"teacher_disagreement_score":0.0034719985,"about_ca_system_score_codex":0.00062684505,"about_ca_system_score_gemma":0.00036225648,"threshold_uncertainty_score":0.018361926},"labels":[],"label_agreement":null},{"id":"W2204352964","doi":"10.1007/978-3-319-24912-4_13","title":"Generating Software Documentation in Use Case Maps from Filtered Execution Traces","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Documentation; Software; Software documentation; Software construction; Software engineering; Notation; Software development; Software system; Programming language; Software framework; Software maintenance; Verification and validation; Data mining","score_opus":0.04242136738001418,"score_gpt":0.28739411155629774,"score_spread":0.24497274417628356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2204352964","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.079609305,0.0004203605,0.83986294,0.0002646785,0.0001480062,0.0004502603,0.004516994,0.06883701,0.005890434],"genre_scores_gemma":[0.32905307,0.00055347505,0.64319795,0.000059565493,0.000050018545,0.0004354332,0.012874235,0.0076859114,0.006090427],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989261,0.00016313573,0.00006896064,0.0002037832,0.0005503197,0.000087702756],"domain_scores_gemma":[0.9936107,0.0032637385,0.0004358031,0.0010442509,0.001496972,0.00014857907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011106695,0.0016545674,0.00097488123,0.005726409,0.0007954941,0.0024247805,0.0014935483,0.0011949098,0.006124282],"category_scores_gemma":[0.012301557,0.0010766176,0.0016366434,0.0033945662,0.00049788837,0.0021607052,0.0019865714,0.0013408349,0.002918783],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000623127,0.00040408038,0.011978348,0.0014786948,0.00019660727,0.002202123,0.0023461517,0.111239314,0.021789016,0.01785997,0.027053485,0.802829],"study_design_scores_gemma":[0.00006835496,0.00013132383,0.0040557515,0.00034889483,0.00016150506,0.00065029954,0.0006172403,0.8930401,0.04333735,0.032704003,0.024794184,0.00009098289],"about_ca_topic_score_codex":0.0064802673,"about_ca_topic_score_gemma":0.008622161,"teacher_disagreement_score":0.0064802673,"about_ca_system_score_codex":0.00081309304,"about_ca_system_score_gemma":0.0020503816,"threshold_uncertainty_score":0.020487785},"labels":[],"label_agreement":null},{"id":"W22103706","doi":"10.1111/j.2044-8287.2011.02037.x","title":"Fuzzy Analogy: A New Approach for Software Cost Estimation","year":2001,"lang":"en","type":"article","venue":"British Journal of Health Psychology","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Analogy; Categorical variable; Computer science; Software; Software sizing; Software development; Fuzzy logic; Data mining; Software quality; Estimation; Software metric; Artificial intelligence; Field (mathematics); Machine learning; Software construction; Systems engineering; Mathematics; Programming language; Engineering","score_opus":0.06792452317148749,"score_gpt":0.39047827204493,"score_spread":0.3225537488734425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W22103706","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00119374,0.00007563677,0.99735886,0.00004597676,0.000017720651,0.000026765629,0.000025810506,0.000097129734,0.0011583104],"genre_scores_gemma":[0.11074435,0.000421289,0.8838223,0.00006463844,0.000097106094,0.00037725226,0.0001279264,0.000098265926,0.0042468295],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99873215,0.0005165681,0.00008070226,0.00018094567,0.0004409136,0.000048793845],"domain_scores_gemma":[0.99770516,0.0016249324,0.00013470693,0.00018716425,0.0003035633,0.00004450521],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019461641,0.0010217986,0.0012671425,0.0031385708,0.0005940268,0.0016686138,0.0017778415,0.0013889618,0.008284176],"category_scores_gemma":[0.010096533,0.0005659331,0.0014026574,0.0027557479,0.0006950414,0.0025046733,0.0014571159,0.001781713,0.0014173088],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007351916,0.00008309546,0.00094947655,0.00021422561,0.00012446182,0.00014332109,0.00012638733,0.48398125,0.0016922642,0.19463512,0.0026907441,0.31528613],"study_design_scores_gemma":[0.000007994511,0.000022417933,0.000106456035,0.00001822518,0.00001272819,0.0000476546,0.00001586126,0.9197816,0.0002516866,0.07723217,0.0024907568,0.000012421704],"about_ca_topic_score_codex":0.0040812474,"about_ca_topic_score_gemma":0.0032635103,"teacher_disagreement_score":0.008284176,"about_ca_system_score_codex":0.0008197864,"about_ca_system_score_gemma":0.0008981687,"threshold_uncertainty_score":0.027713358},"labels":[],"label_agreement":null},{"id":"W2214184477","doi":"10.18438/b8p31c","title":"Multiple Databases are Needed to Search the Journal Literature on Computer Science","year":2015,"lang":"en","type":"article","venue":"Evidence Based Library and Information Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Scopus; Information retrieval; Citation database; Subject (documents); Database; Web of science; Ranking (information retrieval); Bibliographic database; Bibliometrics; Scientometrics; Online database; Library science; World Wide Web; MEDLINE","score_opus":0.04325352244456667,"score_gpt":0.3030756420343228,"score_spread":0.25982211958975615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2214184477","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071682828,0.72399896,0.034366522,0.09584553,0.01293999,0.0051986426,0.037460044,0.0030348366,0.07998731],"genre_scores_gemma":[0.04262952,0.61655146,0.23812856,0.031126086,0.006848536,0.009464279,0.04322001,0.0014991038,0.010532522],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9439037,0.013177585,0.024897022,0.003056422,0.013760742,0.0012045293],"domain_scores_gemma":[0.602324,0.22232343,0.032928374,0.021850988,0.1144618,0.0061113792],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.034059003,0.0023164847,0.009192673,0.1150779,0.004219653,0.02385142,0.0072095864,0.0073907855,0.06469707],"category_scores_gemma":[0.2186534,0.0024721469,0.0048370957,0.14506938,0.0020949317,0.031430718,0.009309247,0.004273734,0.036200088],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023549909,0.00016127035,0.002534034,0.11976141,0.0012819956,0.0006844124,0.0012155477,0.00038062435,0.0011587803,0.01633771,0.14391415,0.7123347],"study_design_scores_gemma":[0.00016652746,0.00011575143,0.004433682,0.12991904,0.0021374333,0.0008434538,0.004917562,0.0005087184,0.00093874615,0.021981161,0.8338187,0.00021923611],"about_ca_topic_score_codex":0.0055341725,"about_ca_topic_score_gemma":0.009256523,"teacher_disagreement_score":0.965941,"about_ca_system_score_codex":0.007857651,"about_ca_system_score_gemma":0.028337153,"threshold_uncertainty_score":0.21643329},"labels":[],"label_agreement":null},{"id":"W2220281200","doi":"10.1109/case.1993.634815","title":"Using virtual subsystems in project management","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software project management; Software engineering; Software system; Project management; Process (computing); Software; Reverse engineering; Systems engineering; Software construction; Engineering; Programming language","score_opus":0.08523613634112569,"score_gpt":0.30211637975797023,"score_spread":0.21688024341684453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2220281200","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11337285,0.0020185304,0.8062685,0.002838656,0.00031305978,0.00014387947,0.00004599765,0.0011363805,0.073862135],"genre_scores_gemma":[0.67590845,0.0014784082,0.309156,0.00037161843,0.00007379151,0.00029366056,0.00008845916,0.00021535416,0.012414195],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9950271,0.0038761005,0.00013371813,0.00024847826,0.00052747776,0.0001871758],"domain_scores_gemma":[0.99569726,0.0017056948,0.00051272695,0.0013652282,0.00040129002,0.0003177564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036456631,0.00046373293,0.00022239577,0.0012352207,0.0015786836,0.0044194316,0.0007260501,0.0008199184,0.002777048],"category_scores_gemma":[0.0074028806,0.00047948712,0.00033694028,0.0011772464,0.003169965,0.0049421364,0.004402403,0.00074664067,0.0009989062],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016393226,0.0001172491,0.007391794,0.00022904496,0.000059961785,0.00059105176,0.013439177,0.025553374,0.005051798,0.5466492,0.009135699,0.39161772],"study_design_scores_gemma":[0.00011885738,0.0010999726,0.0070164595,0.0005823323,0.00022483116,0.0021769998,0.009348025,0.083629794,0.011822128,0.36282447,0.5209568,0.00019923839],"about_ca_topic_score_codex":0.0019617334,"about_ca_topic_score_gemma":0.0022879818,"teacher_disagreement_score":0.0044194316,"about_ca_system_score_codex":0.0013161892,"about_ca_system_score_gemma":0.0016021858,"threshold_uncertainty_score":0.019280314},"labels":[],"label_agreement":null},{"id":"W2225528761","doi":"","title":"Measurement convertibility : from function points to COSMIC-FFP","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Convertibility; Set (abstract data type); Function (biology); Computer science; Economics; Monetary economics","score_opus":0.03526869225338885,"score_gpt":0.25481711166229615,"score_spread":0.2195484194089073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2225528761","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39914915,0.0059483256,0.41284922,0.0014630416,0.00043916414,0.00053411286,0.006336215,0.0035127518,0.16976807],"genre_scores_gemma":[0.90140253,0.0010014159,0.090401694,0.00022075653,0.00010565128,0.00026660462,0.003091755,0.0005723308,0.0029372652],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9794629,0.005429004,0.0013179628,0.0016928966,0.011668784,0.0004284752],"domain_scores_gemma":[0.9391093,0.035598755,0.004377391,0.010038351,0.010514911,0.00036123607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009399718,0.0011724913,0.0005400664,0.012325838,0.00088057434,0.0036690703,0.0016495789,0.00079134863,0.005866251],"category_scores_gemma":[0.08264755,0.00045082238,0.0009183331,0.01310849,0.0020672753,0.0071265753,0.0030009786,0.0018443221,0.001195474],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058244343,0.00013472924,0.14007105,0.0012815051,0.00025924665,0.00059379684,0.0076456517,0.0062495344,0.008786833,0.15453972,0.009583492,0.67027205],"study_design_scores_gemma":[0.00007612334,0.0009579997,0.45490977,0.0016101062,0.00041646225,0.004662546,0.008421763,0.04081964,0.058827356,0.22770664,0.20116049,0.0004310748],"about_ca_topic_score_codex":0.005590694,"about_ca_topic_score_gemma":0.0030181345,"teacher_disagreement_score":0.012325838,"about_ca_system_score_codex":0.0021020153,"about_ca_system_score_gemma":0.00085355225,"threshold_uncertainty_score":0.04971105},"labels":[],"label_agreement":null},{"id":"W2225812146","doi":"10.1007/978-3-642-27207-3_14","title":"COSMIC Functional Size Measurement Using UML Models","year":2011,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Unified Modeling Language; Computer science; COSMIC cancer database; Programming language; Astronomy; Physics; Software","score_opus":0.18651982694562994,"score_gpt":0.2993028095764839,"score_spread":0.11278298263085398,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2225812146","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2816495,0.0011405312,0.6116541,0.0003606111,0.00024460375,0.00009749709,0.0019767163,0.011862324,0.09101411],"genre_scores_gemma":[0.8757354,0.00032860608,0.11508343,0.00008733982,0.000039791586,0.000101609345,0.0021735604,0.00133647,0.005113856],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99758005,0.0005172788,0.00010035417,0.00024391427,0.0014347946,0.00012355186],"domain_scores_gemma":[0.99049133,0.004180235,0.0006939751,0.0023622767,0.0020640134,0.0002082803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016503164,0.00067566155,0.00040326273,0.004814537,0.0005211172,0.001520811,0.001077903,0.00074692845,0.0051911063],"category_scores_gemma":[0.013348365,0.00035492476,0.000655826,0.0030221718,0.0006589266,0.002459682,0.000984279,0.000671377,0.0011986586],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039821732,0.00016702244,0.041456804,0.00049603486,0.00014771833,0.00026312776,0.0013925992,0.09768972,0.033230588,0.2741146,0.01685911,0.53378445],"study_design_scores_gemma":[0.000038279766,0.00032045142,0.040186115,0.0002733756,0.00018743087,0.0009597312,0.0005592318,0.65876496,0.083084516,0.14763573,0.06783447,0.00015582534],"about_ca_topic_score_codex":0.0030149135,"about_ca_topic_score_gemma":0.0029432017,"teacher_disagreement_score":0.0051911063,"about_ca_system_score_codex":0.0011962138,"about_ca_system_score_gemma":0.0006235083,"threshold_uncertainty_score":0.017365992},"labels":[],"label_agreement":null},{"id":"W2226551422","doi":"","title":"Assessing RBFN-Based Software Cost Estimation Models (S).","year":2013,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Estimation; Software; Reliability engineering; Engineering; Systems engineering; Operating system","score_opus":0.020617071679228403,"score_gpt":0.25908039871777366,"score_spread":0.23846332703854525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2226551422","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38746765,0.0023378122,0.6028076,0.0007535284,0.00014286513,0.00016468362,0.0007281718,0.000998004,0.0045997235],"genre_scores_gemma":[0.91702235,0.00031536006,0.08102676,0.000067930574,0.000027951623,0.00007838719,0.0007340828,0.00006794962,0.00065917353],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964678,0.001777526,0.00017448336,0.00033051494,0.0011048231,0.00014472616],"domain_scores_gemma":[0.9764772,0.016688777,0.0013285958,0.0013325675,0.0039143977,0.00025846603],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009014936,0.0008905648,0.00081039383,0.0017140337,0.00028698292,0.001297484,0.0013384133,0.0016415688,0.0010737972],"category_scores_gemma":[0.051915083,0.00044493607,0.0007225041,0.0011343003,0.00031256487,0.001874958,0.00089546543,0.0010673043,0.00035840608],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033025074,0.00009999064,0.007472299,0.000117795185,0.00016972216,0.000027599785,0.00003529713,0.920387,0.0008647153,0.002385307,0.0007702671,0.06733976],"study_design_scores_gemma":[0.0000075476823,0.00006275018,0.001218416,0.000016469532,0.000018679786,0.000015461457,0.000010856149,0.9970362,0.00037272685,0.0010556978,0.00017811487,0.0000071084146],"about_ca_topic_score_codex":0.018618286,"about_ca_topic_score_gemma":0.011085644,"teacher_disagreement_score":0.018618286,"about_ca_system_score_codex":0.0017211017,"about_ca_system_score_gemma":0.0016017992,"threshold_uncertainty_score":0.047676086},"labels":[],"label_agreement":null},{"id":"W2227838144","doi":"","title":"Introduciendo conceptos de metrologia en el diseno de medidas de software","year":2008,"lang":"es","type":"article","venue":"Conferencia Iberoamericana de Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Humanities; Computer science; Philosophy","score_opus":0.014995610958120801,"score_gpt":0.26206525579649886,"score_spread":0.24706964483837807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2227838144","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012385031,0.027173692,0.8843009,0.0068286723,0.0018844861,0.00028358598,0.0009516331,0.0015772867,0.06461476],"genre_scores_gemma":[0.17721538,0.025475701,0.7550165,0.0039577857,0.003270738,0.00087271474,0.0019514235,0.001191756,0.03104801],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99134946,0.0027612285,0.001386887,0.0012884673,0.0029157766,0.00029808385],"domain_scores_gemma":[0.98560613,0.008351473,0.0013224767,0.0017043806,0.0026834006,0.00033214083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066816197,0.001419641,0.0009817414,0.006012646,0.0018964886,0.007828552,0.0019935207,0.0032123683,0.009341321],"category_scores_gemma":[0.014514127,0.0008917179,0.001748617,0.0046927566,0.007377759,0.014139093,0.004140489,0.0058744205,0.00323698],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000807801,0.00003567771,0.0014833991,0.0014935927,0.000047929134,0.00032001978,0.0036016754,0.0014773567,0.0055439644,0.83790326,0.008140454,0.13987194],"study_design_scores_gemma":[0.000026974509,0.00012873815,0.0017336818,0.0018860908,0.00009427464,0.0018190482,0.001505905,0.0064568375,0.0048842016,0.3535733,0.6277663,0.0001246203],"about_ca_topic_score_codex":0.0034000247,"about_ca_topic_score_gemma":0.002156097,"teacher_disagreement_score":0.009341321,"about_ca_system_score_codex":0.0024137339,"about_ca_system_score_gemma":0.003097028,"threshold_uncertainty_score":0.035336196},"labels":[],"label_agreement":null},{"id":"W2233589089","doi":"","title":"Mesure du coût de la qualité logicielle d'un projet d'envergure de la société Bombardier Transport","year":2009,"lang":"fr","type":"article","venue":"Génie logiciel","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Humanities; Political science; Philosophy","score_opus":0.01983929838237332,"score_gpt":0.302063252822164,"score_spread":0.2822239544397907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2233589089","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9108082,0.00048315545,0.030788926,0.0026083505,0.000056480538,0.00044442696,0.0013447031,0.0011160374,0.052349806],"genre_scores_gemma":[0.9301652,0.00040799982,0.034360442,0.00015662106,0.00001997648,0.00040871446,0.001711701,0.00016366321,0.032605767],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99358565,0.0019246797,0.00025257032,0.00063193217,0.0028942353,0.0007110763],"domain_scores_gemma":[0.96383035,0.0087120235,0.002595737,0.0024166452,0.018231252,0.0042139804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010717514,0.0006133868,0.00039167103,0.0029482616,0.0019172381,0.0039524026,0.0010701749,0.0008434541,0.008731609],"category_scores_gemma":[0.020555995,0.00040199692,0.0005356822,0.0027585523,0.001185133,0.001994147,0.0022781326,0.0014567099,0.0016226103],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017680879,0.0017248972,0.28243068,0.0021083253,0.00032476662,0.0008009286,0.06961737,0.020398118,0.039565,0.022573981,0.0187856,0.53990227],"study_design_scores_gemma":[0.00018169034,0.002902696,0.7451511,0.0009439014,0.00028913203,0.00035936185,0.05145965,0.035431955,0.04448924,0.0065492974,0.11196607,0.00027595222],"about_ca_topic_score_codex":0.074761316,"about_ca_topic_score_gemma":0.08966917,"teacher_disagreement_score":0.074761316,"about_ca_system_score_codex":0.007883529,"about_ca_system_score_gemma":0.011288671,"threshold_uncertainty_score":0.14865232},"labels":[],"label_agreement":null},{"id":"W2237502633","doi":"10.1007/s10115-015-0909-5","title":"A feature location approach supported by time-aware weighting of terms associated with developer expertise profiles","year":2016,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Feature (linguistics); Weighting; Context (archaeology); Data mining; Metadata; Feature model; Software; Source code; Information retrieval; Term (time); Set (abstract data type); World Wide Web","score_opus":0.008221255373777811,"score_gpt":0.2118801598688599,"score_spread":0.20365890449508212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2237502633","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10201124,0.001728563,0.8844143,0.0004632368,0.00026856,0.00026783333,0.002422932,0.0049737957,0.0034494854],"genre_scores_gemma":[0.40606517,0.0005524205,0.58184904,0.00013839663,0.00018415693,0.00020209774,0.0034279525,0.0006625889,0.0069182087],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99737656,0.00043381067,0.0002838864,0.00068482285,0.0010195913,0.00020137837],"domain_scores_gemma":[0.992191,0.002064954,0.00096126046,0.0013366733,0.0030360387,0.00040997274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002081121,0.0009599379,0.0010231584,0.008062457,0.00078562804,0.0020720558,0.0015308106,0.0012623619,0.0024482638],"category_scores_gemma":[0.012007576,0.00041422923,0.0009557488,0.0066773137,0.00030411998,0.0033235902,0.0020467688,0.0010789422,0.0017740628],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008364763,0.0004632462,0.015146765,0.00052692444,0.00025812341,0.0002793504,0.0008961269,0.005905837,0.09198401,0.005713934,0.012733757,0.86525536],"study_design_scores_gemma":[0.00023161391,0.0009333374,0.061230958,0.000358679,0.0012155115,0.002458565,0.0016805,0.7100278,0.10815516,0.043126628,0.07013902,0.00044224973],"about_ca_topic_score_codex":0.007100055,"about_ca_topic_score_gemma":0.015903711,"teacher_disagreement_score":0.008062457,"about_ca_system_score_codex":0.0007499826,"about_ca_system_score_gemma":0.0017374958,"threshold_uncertainty_score":0.01411742},"labels":[],"label_agreement":null},{"id":"W22384488","doi":"","title":"Using Structure-Based Recommendations to Facilitate Discoverability in APIs","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Discoverability; Computer science; Task (project management); Application programming interface; World Wide Web; Key (lock); Eclipse; Software engineering; Programming language; Engineering; Operating system","score_opus":0.12579657951128936,"score_gpt":0.34092713602133823,"score_spread":0.21513055651004886,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W22384488","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4831203,0.0016302192,0.47761896,0.0028874131,0.00014459636,0.0011818394,0.0013069467,0.017636165,0.014473624],"genre_scores_gemma":[0.48721513,0.0004398968,0.5064609,0.00021531525,0.000047496258,0.00034158112,0.0014364518,0.0004827489,0.003360496],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965406,0.001411598,0.00024415835,0.0006837358,0.0009675963,0.00015234458],"domain_scores_gemma":[0.9571047,0.032185666,0.0029208292,0.0038196184,0.003322283,0.00064697873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045071794,0.0013234742,0.0008704297,0.0034986292,0.0008905067,0.0021536066,0.0020727196,0.001939326,0.0033904149],"category_scores_gemma":[0.047416124,0.0009612221,0.0006795233,0.0014047815,0.0005059617,0.00493365,0.0016222526,0.001535494,0.0013131347],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009969907,0.0016374566,0.06563217,0.0014777319,0.00024319015,0.0013636444,0.007918382,0.016991915,0.02847385,0.0061789076,0.016571188,0.8525145],"study_design_scores_gemma":[0.0008675179,0.0016995352,0.054683402,0.0011607731,0.0010041517,0.0021133092,0.005299654,0.74623686,0.05412894,0.033855766,0.098427385,0.0005227234],"about_ca_topic_score_codex":0.006020588,"about_ca_topic_score_gemma":0.016350985,"teacher_disagreement_score":0.006020588,"about_ca_system_score_codex":0.00073664857,"about_ca_system_score_gemma":0.0016733197,"threshold_uncertainty_score":0.023836553},"labels":[],"label_agreement":null},{"id":"W2240294276","doi":"10.1109/ase.2015.94","title":"Investigating Program Behavior Using the Texada LTL Specifications Miner","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Program comprehension; Programming language; Key (lock); Source code; Code (set theory); Software engineering; Software; Software system; Operating system","score_opus":0.23323892312488567,"score_gpt":0.3604952784969496,"score_spread":0.12725635537206395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2240294276","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015509514,0.00019741213,0.895008,0.00050928333,0.00005624528,0.0002242047,0.009002834,0.07431636,0.0051760175],"genre_scores_gemma":[0.12671277,0.00038339107,0.8347628,0.0004430992,0.000028643239,0.00085690524,0.019254815,0.011225907,0.0063316813],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973922,0.00070371217,0.0002836397,0.0004944314,0.0010167297,0.00010928455],"domain_scores_gemma":[0.98767823,0.007729695,0.0009809415,0.0020573037,0.0013654038,0.00018841116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034037284,0.0011457839,0.0007042211,0.0027027593,0.0008449958,0.0023806805,0.001938325,0.0011977843,0.010567602],"category_scores_gemma":[0.019932259,0.0012402436,0.0020373662,0.0014070389,0.00117288,0.0038657133,0.0025909264,0.0020043051,0.004167087],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012217067,0.00049215293,0.025898244,0.004098761,0.00045927125,0.0027506559,0.004299378,0.13002278,0.06788322,0.15583785,0.11592556,0.4911103],"study_design_scores_gemma":[0.00017450894,0.00021424492,0.0019116406,0.00034550365,0.00013338543,0.0008534071,0.0006167274,0.70553607,0.061895084,0.087110996,0.14107074,0.00013767573],"about_ca_topic_score_codex":0.0027627386,"about_ca_topic_score_gemma":0.0062048226,"teacher_disagreement_score":0.010567602,"about_ca_system_score_codex":0.00090481655,"about_ca_system_score_gemma":0.0025937953,"threshold_uncertainty_score":0.03535217},"labels":[],"label_agreement":null},{"id":"W2241447802","doi":"10.1109/ase.2015.47","title":"Semantic Slicing of Software Version Histories (T)","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Program slicing; Slicing; Set (abstract data type); Programming language; Correctness; Java; Software engineering; Software; Software maintenance; Feature (linguistics); Software system; World Wide Web","score_opus":0.026768829481107058,"score_gpt":0.25514873475056016,"score_spread":0.2283799052694531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2241447802","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08423538,0.0004915034,0.9067586,0.00022857358,0.000049488397,0.0002655582,0.0013242968,0.004615471,0.0020311908],"genre_scores_gemma":[0.44448173,0.00029382837,0.54951495,0.00009472245,0.000042670657,0.00021093874,0.0030619164,0.0010338168,0.0012654471],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951799,0.0011734256,0.0007157064,0.0010895089,0.0015090506,0.00033238714],"domain_scores_gemma":[0.972284,0.012376461,0.003967616,0.0074192467,0.0032666223,0.0006858894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050323945,0.00086015655,0.0011395264,0.0032706172,0.0008088193,0.0021735395,0.0013759992,0.0008893501,0.0022660312],"category_scores_gemma":[0.026346948,0.0007268805,0.0014474364,0.0024957887,0.0014000848,0.0047175456,0.0019901,0.0010880381,0.0003892857],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010229984,0.00017504193,0.04239523,0.0011729439,0.00025092685,0.0012322782,0.002280267,0.18850198,0.03596086,0.11622812,0.0065677417,0.60421157],"study_design_scores_gemma":[0.00006288602,0.0003597266,0.009873016,0.00028529202,0.0001367504,0.0010727994,0.00061768084,0.7628851,0.06370709,0.142578,0.018304508,0.00011723383],"about_ca_topic_score_codex":0.004778253,"about_ca_topic_score_gemma":0.004695158,"teacher_disagreement_score":0.0050323945,"about_ca_system_score_codex":0.0011450131,"about_ca_system_score_gemma":0.0022057372,"threshold_uncertainty_score":0.02661413},"labels":[],"label_agreement":null},{"id":"W2241839369","doi":"10.11575/prism/31257","title":"Appendix to Information Needs in Bug Reports: Improving Cooperation Between Developers and Users","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Scripting language; World Wide Web; Computer science; Categorization; Computer-supported cooperative work; Replicate; Engineering; Work (physics); Programming language; Artificial intelligence","score_opus":0.010658199773058133,"score_gpt":0.24400025432136804,"score_spread":0.23334205454830992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2241839369","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062783286,0.00021951702,0.014678665,0.0040441584,0.0018393032,0.0050018765,0.8871719,0.0076991906,0.073067084],"genre_scores_gemma":[0.03607987,0.0010372838,0.06251986,0.0023891828,0.00094551645,0.009181577,0.739346,0.0032732596,0.1452275],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981881,0.0004714173,0.00040313287,0.0001959732,0.0006322251,0.000109151064],"domain_scores_gemma":[0.91203,0.042939216,0.004707056,0.0040374147,0.034259163,0.0020271803],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0021387935,0.0009651998,0.00072338805,0.0056112227,0.001109399,0.0015328432,0.0012165525,0.0009115883,0.6064056],"category_scores_gemma":[0.048672356,0.00077610917,0.00047028437,0.0062905275,0.0003177244,0.0015936557,0.001149835,0.00090336497,0.22492649],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008865948,0.00015559705,0.0011599467,0.00045824455,0.000005140947,0.00004554551,0.00008705309,0.00017936868,0.0001803371,0.0006548869,0.9731112,0.0238741],"study_design_scores_gemma":[0.0005109118,0.00028452507,0.022862367,0.00092233776,0.00003162541,0.00030836603,0.00085419725,0.0012249605,0.0018705961,0.0044744723,0.9665728,0.00008280206],"about_ca_topic_score_codex":0.0067048315,"about_ca_topic_score_gemma":0.006022733,"teacher_disagreement_score":0.6064056,"about_ca_system_score_codex":0.0017990121,"about_ca_system_score_gemma":0.002654257,"threshold_uncertainty_score":0.5614146},"labels":[],"label_agreement":null},{"id":"W2245294015","doi":"10.1109/issre.2015.7381799","title":"A similarity-based approach for test case prioritization using historical failure data","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":101,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Regression testing; Computer science; Software quality; Test case; Reliability engineering; Quality assurance; Test (biology); Context (archaeology); Data mining; Code coverage; Software quality assurance; Similarity (geometry); Quality (philosophy); Fault detection and isolation; Test suite; Software; Machine learning; Artificial intelligence; Regression analysis; Software system; Software development; Engineering; Programming language; Software construction","score_opus":0.1918004958848227,"score_gpt":0.33518358440739576,"score_spread":0.14338308852257306,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2245294015","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0626528,0.0005636969,0.93268853,0.00013605705,0.000043817774,0.00056322303,0.0005124442,0.0016311173,0.0012082865],"genre_scores_gemma":[0.5137359,0.00017562885,0.48317316,0.00009042289,0.00008628854,0.000507962,0.0014105043,0.00015004721,0.00067020167],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9869466,0.0030526044,0.0018476176,0.002315746,0.00542678,0.00041060848],"domain_scores_gemma":[0.95892304,0.020370772,0.0066813305,0.003770956,0.009369596,0.0008842967],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066380464,0.0015244473,0.002171185,0.015236154,0.0007684595,0.0019534288,0.002720328,0.0015828294,0.0014573733],"category_scores_gemma":[0.049617,0.0004964134,0.0012954454,0.0068238094,0.0008591343,0.0030119293,0.0018039468,0.0014020195,0.000553161],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006627094,0.00096217525,0.05799434,0.00071648933,0.0006829536,0.00036651714,0.00068157486,0.18359192,0.017675143,0.008112758,0.0027337899,0.72581965],"study_design_scores_gemma":[0.00007086946,0.0007057867,0.015759332,0.00005116033,0.00012096208,0.0005901997,0.00014958234,0.9664121,0.006297652,0.007811318,0.0019409931,0.00009000908],"about_ca_topic_score_codex":0.005663417,"about_ca_topic_score_gemma":0.006542317,"teacher_disagreement_score":0.015236154,"about_ca_system_score_codex":0.0016726579,"about_ca_system_score_gemma":0.0017005071,"threshold_uncertainty_score":0.035105765},"labels":[],"label_agreement":null},{"id":"W2251762750","doi":"10.1007/978-3-319-23727-5_28","title":"The Layered Architecture Recovery as a Quadratic Assignment Problem","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; École de Technologie Supérieure","funders":"","keywords":"Computer science; Heuristic; Architecture; Software architecture; Set (abstract data type); Software; Process (computing); Reference architecture; Software system; Distributed computing; Software engineering; Theoretical computer science; Artificial intelligence; Programming language","score_opus":0.02207220847512984,"score_gpt":0.2611705950751626,"score_spread":0.23909838660003277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251762750","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015315763,0.00030448558,0.97298765,0.0008704262,0.00009761195,0.000051982213,0.000106544474,0.00037858842,0.00988685],"genre_scores_gemma":[0.35279536,0.0009896508,0.6005718,0.0003986333,0.00029618823,0.0002732661,0.00069561874,0.00074076693,0.04323869],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9989716,0.00031176815,0.00004632862,0.00022806355,0.0002638423,0.0001784339],"domain_scores_gemma":[0.9981604,0.0010296326,0.00011157754,0.00036424378,0.00021202066,0.00012201802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013317626,0.000878347,0.0011655915,0.000602566,0.0008424268,0.0023339533,0.002961988,0.0021474508,0.0144744925],"category_scores_gemma":[0.006558543,0.00086592534,0.0012061261,0.0013204128,0.0015270232,0.0046886313,0.0034790756,0.004097753,0.0017242425],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027766143,0.00016572207,0.0004336733,0.0004482639,0.000071177514,0.00020694992,0.00025352972,0.34615844,0.003994136,0.43067643,0.025499282,0.19181469],"study_design_scores_gemma":[0.000036390495,0.000042265277,0.00013298859,0.000035874727,0.000019341487,0.000093252456,0.00007701716,0.64411104,0.00095607154,0.35022393,0.0042540603,0.000017783852],"about_ca_topic_score_codex":0.0026284468,"about_ca_topic_score_gemma":0.0023294794,"teacher_disagreement_score":0.0144744925,"about_ca_system_score_codex":0.0011173759,"about_ca_system_score_gemma":0.0013170782,"threshold_uncertainty_score":0.04842198},"labels":[],"label_agreement":null},{"id":"W2252011085","doi":"","title":"Harmonizing Systems and Software Cost Estimation","year":2009,"lang":"en","type":"article","venue":"DSpace@MIT (Massachusetts Institute of Technology)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"COCOMO; Computer science; Software system; Software development; Estimation; Cost estimate; Software; Software sizing; Software metric; Software engineering; Software construction; Systems engineering; Engineering","score_opus":0.024997370779565126,"score_gpt":0.26148211857958725,"score_spread":0.23648474780002213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252011085","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029105404,0.0004961676,0.9640406,0.00040644634,0.000034146582,0.00006797999,0.000071595365,0.00034775576,0.005429822],"genre_scores_gemma":[0.71336097,0.00086243305,0.28280076,0.0001488294,0.00007305751,0.00022011776,0.00047833807,0.0003250376,0.0017304708],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9803145,0.008751779,0.0009362127,0.0014182795,0.007873194,0.0007061027],"domain_scores_gemma":[0.98238945,0.010220762,0.0014879477,0.0026547324,0.003119659,0.0001273023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01129459,0.0011801724,0.0015523244,0.0047405674,0.00047198505,0.0039689997,0.0018155505,0.0011434659,0.0019131955],"category_scores_gemma":[0.050409462,0.000942669,0.00088899763,0.004901194,0.0009157721,0.0054236525,0.0036952468,0.0014602497,0.00056091975],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008722178,0.000099249024,0.0064578108,0.0002000819,0.00013964943,0.000081895014,0.00039843473,0.61824334,0.0012912414,0.085802324,0.0012635469,0.28593528],"study_design_scores_gemma":[0.000014245383,0.00012033026,0.005209245,0.00012308222,0.000050036175,0.00007847143,0.00037364918,0.9090346,0.0026713917,0.075149514,0.007126558,0.0000488461],"about_ca_topic_score_codex":0.0044767195,"about_ca_topic_score_gemma":0.0036419425,"teacher_disagreement_score":0.01129459,"about_ca_system_score_codex":0.0016921769,"about_ca_system_score_gemma":0.0022801731,"threshold_uncertainty_score":0.0597322},"labels":[],"label_agreement":null},{"id":"W2252116325","doi":"","title":"Reverse engineering of object-oriented code into Umple using an incremental and rule-based approach","year":2014,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Programming language; Unified Modeling Language; Java; Reverse engineering; Code (set theory); Object-oriented programming; Code generation; Source code; Metamodeling; KPI-driven code analysis; Software engineering; Theoretical computer science; Static program analysis; Software development; Software; Operating system","score_opus":0.015916191995144578,"score_gpt":0.2391135352928475,"score_spread":0.22319734329770294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252116325","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010527516,0.00006338486,0.98162675,0.00011706925,0.00005771114,0.0001920625,0.00004752938,0.004936237,0.0024317047],"genre_scores_gemma":[0.070649534,0.00015302753,0.9240435,0.00024296703,0.000024014567,0.00017743584,0.00033711232,0.0011147828,0.0032576353],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9960026,0.0010851341,0.00029808958,0.0006197598,0.0017891995,0.00020519661],"domain_scores_gemma":[0.9838125,0.0061855544,0.0010559977,0.006595719,0.002156887,0.00019341963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003808532,0.00078670355,0.0005459826,0.0012766528,0.00065578823,0.002162825,0.0022943376,0.0011326218,0.0015255488],"category_scores_gemma":[0.017406922,0.0006461645,0.0012232146,0.0004841681,0.0015905298,0.0021356153,0.0022917085,0.0022749775,0.0010471276],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003376468,0.0005975172,0.0056100697,0.00082367455,0.00021840777,0.002651187,0.0022641432,0.074273765,0.108461566,0.12112178,0.0056866673,0.67795354],"study_design_scores_gemma":[0.000082996936,0.0004304488,0.001738886,0.0004550625,0.00023462204,0.0021740138,0.00038378802,0.48704642,0.31421825,0.06100434,0.13207668,0.00015447126],"about_ca_topic_score_codex":0.0010808613,"about_ca_topic_score_gemma":0.0018532578,"teacher_disagreement_score":0.003808532,"about_ca_system_score_codex":0.0005425404,"about_ca_system_score_gemma":0.0015434478,"threshold_uncertainty_score":0.020141661},"labels":[],"label_agreement":null},{"id":"W2253985078","doi":"","title":"A multi-dimensional approach for analyzing software artifacts","year":2013,"lang":"en","type":"article","venue":"Espace ÉTS (ETS)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Software; Programming language","score_opus":0.026671103101190984,"score_gpt":0.26931332692100257,"score_spread":0.2426422238198116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2253985078","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013647589,0.0002587328,0.9832312,0.000121422076,0.00002346411,0.00010306904,0.00035875462,0.00033756674,0.0019181621],"genre_scores_gemma":[0.120329954,0.00024431042,0.8769849,0.000055892204,0.000026129332,0.00031521718,0.00066850375,0.00005063014,0.0013244505],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99728644,0.0007238966,0.00029190915,0.00053691334,0.0010013112,0.00015959311],"domain_scores_gemma":[0.99716043,0.0011182271,0.00043895477,0.00046902752,0.00063871196,0.00017461098],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013823345,0.00094837183,0.00078902417,0.008035092,0.0011901412,0.0040264,0.0009693344,0.00083140534,0.0023658336],"category_scores_gemma":[0.004492057,0.00052388,0.002199435,0.0051390626,0.0010611705,0.0020309533,0.002684692,0.0011103177,0.00057676155],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029783003,0.0005605561,0.03233171,0.00096937694,0.0006041879,0.00065520086,0.0037329248,0.06456429,0.042103227,0.11790247,0.005840326,0.7304379],"study_design_scores_gemma":[0.000023392024,0.00023905344,0.02099238,0.00023673078,0.0001977786,0.00080296677,0.0027910133,0.8299674,0.0111398315,0.114220455,0.019224782,0.00016412645],"about_ca_topic_score_codex":0.0018946967,"about_ca_topic_score_gemma":0.003122403,"teacher_disagreement_score":0.008035092,"about_ca_system_score_codex":0.00092311995,"about_ca_system_score_gemma":0.0010961766,"threshold_uncertainty_score":0.007914484},"labels":[],"label_agreement":null},{"id":"W2258726425","doi":"10.1145/2901739.2901766","title":"The unreasonable effectiveness of traditional information retrieval in crash report deduplication","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Mozilla Foundation; Mitacs","keywords":"Crash; Computer science; Data deduplication; Software; Scalability; Information retrieval; Precision and recall; Database; Data science; Software engineering; World Wide Web; Operating system","score_opus":0.015265970731580713,"score_gpt":0.2492900997646688,"score_spread":0.23402412903308809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2258726425","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5347606,0.038390085,0.35976186,0.006344651,0.0018894018,0.0014904386,0.005781776,0.032156546,0.019424595],"genre_scores_gemma":[0.680832,0.006368766,0.29771826,0.001216283,0.00056444807,0.0002869801,0.005545086,0.000920068,0.0065481183],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9899833,0.0027312539,0.0012383248,0.0017106459,0.00384984,0.000486535],"domain_scores_gemma":[0.94151884,0.033096995,0.0027848922,0.01617038,0.005931886,0.00049703417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011945706,0.0013889173,0.0019201923,0.0058975616,0.0018473134,0.004377502,0.003188596,0.0020830382,0.0017056398],"category_scores_gemma":[0.052897815,0.000698969,0.000933059,0.0069784517,0.0017103308,0.010578661,0.0021049678,0.0016967398,0.0034802656],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001385154,0.00057595794,0.014519702,0.002554147,0.00043543053,0.00044141445,0.0010384341,0.024439946,0.040689886,0.0046442645,0.046350237,0.86292547],"study_design_scores_gemma":[0.00076023187,0.0038632327,0.045277573,0.0010224781,0.0011895106,0.008083501,0.005336393,0.3945798,0.3711771,0.037376154,0.13053645,0.00079758884],"about_ca_topic_score_codex":0.004847452,"about_ca_topic_score_gemma":0.006315708,"teacher_disagreement_score":0.011945706,"about_ca_system_score_codex":0.0013976502,"about_ca_system_score_gemma":0.0022619017,"threshold_uncertainty_score":0.06317568},"labels":[],"label_agreement":null},{"id":"W2258866920","doi":"","title":"Leveraging Software Clones for Software Comprehension: Techniques and Practice","year":2015,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Humanities; Philosophy","score_opus":0.03283854412911973,"score_gpt":0.28275347689457997,"score_spread":0.24991493276546023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2258866920","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03815358,0.0032404,0.94431394,0.0012039643,0.00006734292,0.0001765695,0.00011355825,0.00801405,0.004716548],"genre_scores_gemma":[0.17632735,0.002496279,0.8143594,0.00027265304,0.00008091508,0.00014214711,0.00039568407,0.0016871365,0.004238426],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.991478,0.0023545215,0.0005746358,0.0015656505,0.003644161,0.00038298182],"domain_scores_gemma":[0.9738304,0.013730601,0.0019508876,0.004377894,0.0057148756,0.0003953649],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050326143,0.0014730077,0.00090888323,0.0049123513,0.0009225058,0.0040778727,0.0023778363,0.0023481327,0.0041875737],"category_scores_gemma":[0.030944543,0.0011364043,0.001276437,0.003048925,0.002048981,0.0066395276,0.00202641,0.0020543104,0.0023820535],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018267678,0.00019070871,0.010651217,0.0008596879,0.00010871895,0.00053371757,0.005236079,0.008471987,0.060414314,0.02041585,0.0046941596,0.8882409],"study_design_scores_gemma":[0.00020423801,0.0010375481,0.017353792,0.0025514334,0.00051220524,0.00843324,0.0049908594,0.3309903,0.31562907,0.095235705,0.22257878,0.00048282254],"about_ca_topic_score_codex":0.0037431668,"about_ca_topic_score_gemma":0.0035080763,"teacher_disagreement_score":0.0050326143,"about_ca_system_score_codex":0.0013883058,"about_ca_system_score_gemma":0.0019890498,"threshold_uncertainty_score":0.026615322},"labels":[],"label_agreement":null},{"id":"W2259658967","doi":"","title":"Analysis of the ISO 9126 on software product quality evaluation from the metrology and ISO 15939 perspectives","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Metrology; Software; Software measurement; Documentation; Systems engineering; Verification and validation; Software quality; Engineering; Quality assurance; Software engineering; Computer science; Reliability engineering; Software development; Mathematics; Operations management; Statistics; Programming language","score_opus":0.03141104115357844,"score_gpt":0.3229956345772131,"score_spread":0.29158459342363463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2259658967","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3465209,0.022871753,0.3583067,0.010914628,0.001114705,0.0033564977,0.0016968035,0.0014188903,0.25379914],"genre_scores_gemma":[0.7822004,0.0066071963,0.19264577,0.0010581397,0.0004992853,0.0020192266,0.004697646,0.00056749897,0.009704787],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.860352,0.030226462,0.009323809,0.0026169824,0.095344335,0.0021364577],"domain_scores_gemma":[0.89883024,0.03094184,0.0082032895,0.0040624132,0.056890313,0.0010719722],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06364653,0.0018032703,0.0013181983,0.025429755,0.001657811,0.009203298,0.0017493032,0.0018998095,0.0022912922],"category_scores_gemma":[0.08818973,0.00083866843,0.0016338985,0.018169936,0.003813662,0.0055370913,0.003034289,0.0024618898,0.0011362636],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061176025,0.0008345683,0.04320104,0.003323364,0.0002510099,0.0008579713,0.008280748,0.013652584,0.017593412,0.26846114,0.024203066,0.61872935],"study_design_scores_gemma":[0.00028328676,0.00457165,0.2805683,0.0075751273,0.0007225173,0.0031127152,0.017334247,0.08412983,0.033093058,0.091240324,0.47681558,0.00055333273],"about_ca_topic_score_codex":0.007854054,"about_ca_topic_score_gemma":0.0060101496,"teacher_disagreement_score":0.06364653,"about_ca_system_score_codex":0.009071165,"about_ca_system_score_gemma":0.0144841485,"threshold_uncertainty_score":0.3365991},"labels":[],"label_agreement":null},{"id":"W2260251867","doi":"","title":"A design-rule-based constructive approach to building traceable software","year":2009,"lang":"en","type":"dissertation","venue":"TSpace","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Traceability; Requirements traceability; Computer science; Software engineering; Maintainability; Software system; TRACE (psycholinguistics); Software development; Constructive; Process (computing); Systems engineering; Software; Engineering; Programming language; Requirement","score_opus":0.029168824425162915,"score_gpt":0.3184090092066887,"score_spread":0.2892401847815258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2260251867","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00080746727,0.000036197507,0.9975533,0.00010960774,0.000010228326,0.00010497045,0.000015742662,0.0002516557,0.0011108628],"genre_scores_gemma":[0.012764492,0.00008774002,0.9860345,0.00006723612,0.00000989928,0.00019306055,0.00006443345,0.00005105916,0.00072755275],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9904819,0.0033085258,0.0007754787,0.0010285118,0.0041526915,0.00025282564],"domain_scores_gemma":[0.98146147,0.009667065,0.0009831118,0.004291299,0.0033528446,0.00024413932],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009254901,0.0012576358,0.0009987829,0.0022560852,0.0010968398,0.003385517,0.0046965373,0.0018415808,0.0026996352],"category_scores_gemma":[0.021413187,0.0012053766,0.0023972942,0.0015598885,0.0042810026,0.0031800203,0.002464241,0.0039571133,0.0012612224],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007509565,0.0005282166,0.0014213339,0.0011374679,0.00021864641,0.0007735574,0.0019939798,0.15215807,0.027754515,0.45259878,0.003769998,0.3575704],"study_design_scores_gemma":[0.00013806668,0.00043649998,0.00037086738,0.0004572025,0.0002920996,0.00093252136,0.00038353814,0.5769155,0.051907104,0.29686043,0.07118119,0.00012494465],"about_ca_topic_score_codex":0.002399842,"about_ca_topic_score_gemma":0.0033282896,"teacher_disagreement_score":0.009254901,"about_ca_system_score_codex":0.0018343455,"about_ca_system_score_gemma":0.004588113,"threshold_uncertainty_score":0.04894519},"labels":[],"label_agreement":null},{"id":"W2260739911","doi":"10.5539/cis.v9n1p101","title":"A Multistep Approach for Managing the Risks of Software Requirements Volatility","year":2016,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Applied Science Private University","keywords":"Computer science; Business requirements; Software engineering; Software; Requirements management; Risk analysis (engineering); Software requirements; Requirement prioritization; Software development; Volatility (finance); Requirement; Requirements analysis; Business process; Software construction; Operations management; Work in process; Business; Operating system; Engineering","score_opus":0.0705190611327147,"score_gpt":0.32554916306588005,"score_spread":0.25503010193316533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2260739911","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019727722,0.00025694625,0.96273386,0.0020542226,0.00007361682,0.0013292822,0.00007165177,0.001192232,0.01256048],"genre_scores_gemma":[0.27539882,0.00036702232,0.71527034,0.00045074607,0.00004629751,0.00083169073,0.00018747347,0.00010712459,0.0073405057],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.983374,0.005923641,0.001232205,0.0025309282,0.0059327516,0.0010065287],"domain_scores_gemma":[0.98507684,0.0058577037,0.002122137,0.0021488338,0.0036179984,0.0011764597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013580925,0.0017668869,0.001139039,0.004909178,0.0033001127,0.0069860653,0.0042507504,0.0031635375,0.0043146797],"category_scores_gemma":[0.017814465,0.001359979,0.0033887252,0.0020701243,0.0026219408,0.0077781077,0.009202615,0.0044496194,0.0012814697],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051328016,0.001757107,0.032705877,0.0012118177,0.00089421356,0.0035134668,0.012139074,0.12749061,0.028761737,0.2311985,0.0073915194,0.5524227],"study_design_scores_gemma":[0.00013838857,0.0013522517,0.0082109235,0.00075061934,0.0006141401,0.0030812833,0.006750306,0.7114673,0.014206863,0.21240434,0.04053529,0.00048824263],"about_ca_topic_score_codex":0.0050682365,"about_ca_topic_score_gemma":0.0068064034,"teacher_disagreement_score":0.013580925,"about_ca_system_score_codex":0.0034510987,"about_ca_system_score_gemma":0.010421927,"threshold_uncertainty_score":0.0718236},"labels":[],"label_agreement":null},{"id":"W2263084026","doi":"10.11575/prism/22011","title":"A software engineering measurement expert system","year":2003,"lang":"en","type":"dissertation","venue":"PRISM (University of Calgary)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software engineering; Expert system; Systems engineering; Software system; Computer science; Engineering; Software; Operating system; Artificial intelligence","score_opus":0.014363307728865366,"score_gpt":0.20120882853689967,"score_spread":0.1868455208080343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2263084026","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077334004,0.00022693728,0.93683636,0.0008466661,0.00014552435,0.00056473026,0.0008114634,0.030871611,0.021963267],"genre_scores_gemma":[0.08881806,0.00033468602,0.8884221,0.0005822792,0.00010909845,0.0007262137,0.0033144124,0.0013145367,0.016378656],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9945747,0.0011987378,0.00060622074,0.0010999381,0.0023495106,0.00017086526],"domain_scores_gemma":[0.98547715,0.004668537,0.0006943688,0.0022806257,0.0064654835,0.00041378167],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00667272,0.0006722115,0.0009800354,0.0030533762,0.00066807074,0.0025908658,0.0023627102,0.001240011,0.0137298275],"category_scores_gemma":[0.025480842,0.0006233945,0.0005731322,0.0015812658,0.0004190824,0.003513917,0.0021137432,0.0020873088,0.0087264925],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018152955,0.000430893,0.0027441245,0.00048029431,0.00009076621,0.0002496434,0.0008123302,0.012372378,0.017849702,0.030428713,0.045880698,0.888479],"study_design_scores_gemma":[0.00033410973,0.00046143457,0.006251634,0.00065279467,0.00025054967,0.0011970135,0.0006041672,0.42238337,0.044135727,0.05919009,0.46430346,0.00023562275],"about_ca_topic_score_codex":0.0010086977,"about_ca_topic_score_gemma":0.0011042346,"teacher_disagreement_score":0.0137298275,"about_ca_system_score_codex":0.0009257473,"about_ca_system_score_gemma":0.0026701754,"threshold_uncertainty_score":0.045930922},"labels":[],"label_agreement":null},{"id":"W2264470993","doi":"","title":"Functional Complexity Measurement","year":2001,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Complexity management; Worst-case complexity; Programming complexity; Descriptive complexity theory; Computational complexity theory; Component (thermodynamics); Terminology; Theoretical computer science; Software; Software system; Algorithm","score_opus":0.16548067360039131,"score_gpt":0.28271028151297,"score_spread":0.11722960791257869,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2264470993","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07777653,0.00077527587,0.8258002,0.0011562102,0.00023154657,0.00033631787,0.0007553885,0.00071957044,0.09244905],"genre_scores_gemma":[0.87427855,0.00034946096,0.11829325,0.00018032733,0.00017213271,0.00056620693,0.0005915672,0.00012546722,0.005443018],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99216145,0.001852419,0.00047251032,0.0010204768,0.003936336,0.0005568443],"domain_scores_gemma":[0.979806,0.008146094,0.0017354732,0.0037022843,0.0059574125,0.0006527713],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004431467,0.0009906957,0.0005672335,0.0049932715,0.0014916095,0.003506513,0.001547368,0.0012972758,0.007587991],"category_scores_gemma":[0.033625994,0.00030002568,0.0011407341,0.0032384847,0.003235053,0.008151228,0.0023993407,0.0017965698,0.0011008447],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000068078916,0.000092296315,0.009073996,0.00025080485,0.00006589795,0.00008741257,0.0007654024,0.009076606,0.006161811,0.8430407,0.002518476,0.12879854],"study_design_scores_gemma":[0.000015803702,0.00037737007,0.025589805,0.00016358937,0.00006743668,0.00061885023,0.00087178213,0.07704411,0.012984215,0.8566981,0.025439212,0.00012967558],"about_ca_topic_score_codex":0.001793324,"about_ca_topic_score_gemma":0.00079207064,"teacher_disagreement_score":0.007587991,"about_ca_system_score_codex":0.002249804,"about_ca_system_score_gemma":0.0013303289,"threshold_uncertainty_score":0.025384367},"labels":[],"label_agreement":null},{"id":"W2267083253","doi":"10.48550/arxiv.1508.00032","title":"A Neuro-Fuzzy Model with SEER-SEM for Software Effort Estimation","year":2015,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Estimation; Software; Fuzzy logic; Software sizing; Data mining; Software system; Machine learning; Artificial intelligence; Software construction; Systems engineering; Engineering","score_opus":0.08030853294532818,"score_gpt":0.21326078826300315,"score_spread":0.13295225531767496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2267083253","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044347815,0.00031413772,0.9504672,0.00035531112,0.000050710958,0.00011826046,0.00013564681,0.000254995,0.003955929],"genre_scores_gemma":[0.83220094,0.0003399774,0.16393024,0.00011977543,0.000043776094,0.0003148496,0.00015942263,0.000019154148,0.0028719062],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986261,0.0006239002,0.00006471737,0.00031543543,0.00028863837,0.00008112285],"domain_scores_gemma":[0.99813604,0.0010940728,0.00018873204,0.00010410634,0.0004381239,0.000038968905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028087571,0.0008177314,0.000864364,0.0010177274,0.000551814,0.001290442,0.0016979785,0.0014882945,0.0017694109],"category_scores_gemma":[0.0058996985,0.0004043564,0.0009178686,0.0011730755,0.00060376665,0.0012564803,0.0007818912,0.0013131145,0.00040719757],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011409265,0.000111254936,0.002410409,0.0000905969,0.000100472505,0.000107374915,0.00014164545,0.9378064,0.0008355069,0.015664509,0.0005825176,0.04203527],"study_design_scores_gemma":[0.0000038575386,0.000025692812,0.0002712226,0.0000065866407,0.000010429041,0.000010822625,0.0000098235505,0.9971322,0.00013460197,0.0022291203,0.0001584209,0.0000072445996],"about_ca_topic_score_codex":0.012224267,"about_ca_topic_score_gemma":0.013754444,"teacher_disagreement_score":0.012224267,"about_ca_system_score_codex":0.0013067363,"about_ca_system_score_gemma":0.0011645373,"threshold_uncertainty_score":0.024306178},"labels":[],"label_agreement":null},{"id":"W2271119194","doi":"10.11575/prism/31125","title":"Using Structural Generalization to Discover Replacement Functionality for API Evolution","year":2014,"lang":"en","type":"article","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Application programming interface; Software; Generalization; Java; Matching (statistics); Software engineering; Programming language","score_opus":0.07569428380508372,"score_gpt":0.361097115499958,"score_spread":0.28540283169487424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2271119194","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5493462,0.0023069812,0.42298195,0.0007763289,0.00009149942,0.0003812958,0.002557027,0.017057726,0.0045011183],"genre_scores_gemma":[0.7212012,0.00045445148,0.2711662,0.0002472969,0.000051561234,0.00013189911,0.004675594,0.00055264425,0.0015191828],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99766123,0.00047210694,0.00015985145,0.00094801973,0.0006010721,0.0001577377],"domain_scores_gemma":[0.9889703,0.0038337577,0.002129655,0.003247714,0.0015445352,0.00027407907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020879514,0.0011441681,0.0010711674,0.008361125,0.0011378623,0.0011764577,0.0018859806,0.001628534,0.0013715887],"category_scores_gemma":[0.015966829,0.00066290365,0.001445797,0.005415073,0.0008279486,0.003224859,0.0012513727,0.0013444461,0.0008560699],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039314813,0.0004251884,0.28864217,0.00052851223,0.00040478443,0.00041561824,0.0011349543,0.030337071,0.026281796,0.004807519,0.010521582,0.6361077],"study_design_scores_gemma":[0.000059071477,0.00027874496,0.048627086,0.00008527387,0.00030251997,0.00080521795,0.00043220713,0.91019595,0.014914595,0.014053608,0.010163729,0.000081993785],"about_ca_topic_score_codex":0.0141739175,"about_ca_topic_score_gemma":0.032600466,"teacher_disagreement_score":0.0141739175,"about_ca_system_score_codex":0.0010518398,"about_ca_system_score_gemma":0.00175363,"threshold_uncertainty_score":0.028182864},"labels":[],"label_agreement":null},{"id":"W2271947231","doi":"10.7939/r3-wbr5-yd60","title":"Involvement, Contribution and Influence in Github and Stack Overflow","year":2014,"lang":"en","type":"article","venue":"University of Alberta Library","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Reputation; Context (archaeology); Stack (abstract data type); Set (abstract data type); Software; World Wide Web; Code (set theory); Social network (sociolinguistics); Core (optical fiber); Source code; Data science; Social media; Programming language; Telecommunications","score_opus":0.005068555224338849,"score_gpt":0.17605873103240977,"score_spread":0.1709901758080709,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2271947231","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9948397,0.00011318943,0.00026734063,0.00019038166,0.000010479723,0.000017799124,0.000081948776,0.00002538531,0.004453819],"genre_scores_gemma":[0.99719834,0.00010941889,0.0005731451,0.000027081165,0.000029770623,0.000031324056,0.00014779647,0.000022288406,0.0018608352],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99709356,0.0012647265,0.00015691418,0.00033456847,0.00068596273,0.0004642789],"domain_scores_gemma":[0.96180826,0.020326328,0.008773208,0.0013793679,0.0024593268,0.0052535133],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0037796465,0.00029885347,0.00028754465,0.006274878,0.0018085712,0.0036265496,0.0007556799,0.00071777974,0.0019568335],"category_scores_gemma":[0.024629854,0.00026806875,0.00019749643,0.0044802614,0.0019133777,0.002846961,0.0039305133,0.00060580025,0.0004159858],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036885942,0.00020611001,0.7532131,0.00019757557,0.000049300994,0.0013765442,0.174198,0.0004377411,0.003397575,0.0023416856,0.0022887595,0.061924767],"study_design_scores_gemma":[0.000009853275,0.00010864552,0.90536404,0.000081582315,0.00003200628,0.00048671587,0.07853417,0.0016118644,0.0008421011,0.00077779044,0.012103398,0.00004791999],"about_ca_topic_score_codex":0.015087628,"about_ca_topic_score_gemma":0.03254908,"teacher_disagreement_score":0.9937251,"about_ca_system_score_codex":0.0014420198,"about_ca_system_score_gemma":0.0012231983,"threshold_uncertainty_score":0.029999614},"labels":[],"label_agreement":null},{"id":"W2274972169","doi":"","title":"Impacts and detection of design smells","year":2012,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Code smell; Computer science; Principal (computer security); Software engineering; Software design; Software design pattern; Program comprehension; Comprehension; Object-oriented design; Empirical research; Coding (social sciences); Software; Software development; Engineering; Software quality; Software system; Computer security; Programming language","score_opus":0.02286342656875724,"score_gpt":0.28307299815019477,"score_spread":0.26020957158143754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2274972169","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9963057,0.00016401145,0.0019119865,0.000084686086,0.00001131884,0.00008443045,0.00010252207,0.00014330553,0.0011920998],"genre_scores_gemma":[0.9963973,0.00006948713,0.0027832107,0.000052780655,0.000012696359,0.000075321244,0.0002023795,0.000033196215,0.00037365482],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98383236,0.0052670487,0.0014755492,0.0023542806,0.006497217,0.0005735762],"domain_scores_gemma":[0.6739042,0.24322051,0.05324532,0.011172822,0.014163025,0.0042941887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009729785,0.0010804875,0.00068879215,0.0026706376,0.0004623809,0.0022539492,0.00083013286,0.0011801319,0.0019954504],"category_scores_gemma":[0.14106852,0.00062924833,0.0007874333,0.0010644621,0.0009489164,0.0025326286,0.0018385383,0.0013658187,0.00048767298],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024689438,0.0022480953,0.8518478,0.00057102466,0.00035538548,0.0004309046,0.006310708,0.002623474,0.037886105,0.0002401298,0.00065136736,0.09436606],"study_design_scores_gemma":[0.00007786264,0.0030070283,0.9769071,0.000060593527,0.00014356876,0.00036237625,0.0013572862,0.008446173,0.008380395,0.00040642938,0.0007678129,0.00008345884],"about_ca_topic_score_codex":0.001073249,"about_ca_topic_score_gemma":0.0010810194,"teacher_disagreement_score":0.009729785,"about_ca_system_score_codex":0.0006756529,"about_ca_system_score_gemma":0.00064211705,"threshold_uncertainty_score":0.05145663},"labels":[],"label_agreement":null},{"id":"W2277373197","doi":"10.4230/lipics.ecoop.2015.321","title":"Hybrid DOM-Sensitive Change Impact Analysis for JavaScript","year":2015,"lang":"en","type":"article","venue":"DROPS (Schloss Dagstuhl – Leibniz Center for Informatics)","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"JavaScript; Computer science; Asynchronous communication; Static analysis; Unobtrusive JavaScript; Task (project management); Set (abstract data type); Ranking (information retrieval); Distributed computing; Data mining; Artificial intelligence; Programming language; Rich Internet application; Computer network","score_opus":0.05926056696419222,"score_gpt":0.3147781806116966,"score_spread":0.25551761364750436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2277373197","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16769786,0.0006656379,0.81128436,0.00013160506,0.000057155165,0.00042318288,0.0008919651,0.015965056,0.0028831537],"genre_scores_gemma":[0.6730058,0.00017526484,0.32317802,0.00007653781,0.000043236654,0.00019137021,0.0015149381,0.00064904295,0.0011657693],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949988,0.0010710177,0.00037109762,0.0007280527,0.00251417,0.0003168145],"domain_scores_gemma":[0.9891431,0.0048500444,0.0012325207,0.0016987209,0.0027388055,0.00033675274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025843338,0.0011395209,0.0009025902,0.007613788,0.0006473356,0.0016789082,0.001129756,0.0007755783,0.0009542256],"category_scores_gemma":[0.011942477,0.00040558595,0.0009575397,0.0028113506,0.00050058606,0.0018784147,0.0017931466,0.0009985118,0.0005259534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072459923,0.00093394663,0.07915751,0.0006072295,0.0003522804,0.00076161535,0.00082305656,0.0820094,0.0903322,0.006480772,0.005639836,0.7321776],"study_design_scores_gemma":[0.000029500661,0.00023818623,0.026928116,0.000040165192,0.00009445381,0.000465118,0.00021891996,0.9212724,0.040158693,0.006660903,0.0038121033,0.000081503356],"about_ca_topic_score_codex":0.0026888894,"about_ca_topic_score_gemma":0.0040752087,"teacher_disagreement_score":0.007613788,"about_ca_system_score_codex":0.00072128675,"about_ca_system_score_gemma":0.0008544319,"threshold_uncertainty_score":0.013667464},"labels":[],"label_agreement":null},{"id":"W2278479890","doi":"","title":"Software clustering based on behavioural features","year":2007,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Cluster analysis; Computer science; Software; Data mining; Software system; Focus (optics); Software maintenance; Software construction; Software metric; Task (project management); Software analytics; Software sizing; Software engineering; Machine learning; Programming language; Systems engineering; Engineering","score_opus":0.03632589238399748,"score_gpt":0.2942047225007634,"score_spread":0.25787883011676593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2278479890","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3109375,0.0002779951,0.6778143,0.00021898605,0.000033310356,0.0004199805,0.00028805053,0.0013240037,0.008685939],"genre_scores_gemma":[0.8648912,0.000099523284,0.1325315,0.000025311228,0.00002159221,0.00018957444,0.0005074441,0.00013244731,0.0016015279],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979386,0.0003590244,0.00014275462,0.00045175265,0.00091396726,0.00019392293],"domain_scores_gemma":[0.99026185,0.0033096746,0.0015755752,0.0012176043,0.0032721795,0.00036310754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010589154,0.0005842692,0.00058103295,0.0073583825,0.0010712041,0.0016785981,0.001123084,0.00084139674,0.0014847548],"category_scores_gemma":[0.010937354,0.00033052213,0.0009770716,0.0035552566,0.0012153977,0.001935436,0.0011924924,0.0005877362,0.0005694434],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078272447,0.00043962256,0.115551755,0.0008244987,0.00043619325,0.000728736,0.0032663993,0.14973678,0.06596186,0.06827554,0.003220789,0.5907752],"study_design_scores_gemma":[0.000051221217,0.00050181727,0.091866665,0.00015632402,0.00029121348,0.0014188356,0.0013807887,0.80880654,0.026069174,0.062080707,0.007179815,0.00019683265],"about_ca_topic_score_codex":0.0033363523,"about_ca_topic_score_gemma":0.0037759766,"teacher_disagreement_score":0.0073583825,"about_ca_system_score_codex":0.0013066155,"about_ca_system_score_gemma":0.00091850816,"threshold_uncertainty_score":0.009480178},"labels":[],"label_agreement":null},{"id":"W2279884163","doi":"10.1007/978-3-319-24704-5_12","title":"Relational Mathematics for Relative Correctness","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Correctness; Computer science; Theoretical computer science; Relation (database); Algebraic number; Programming language; Program analysis; Algorithm; Algebra over a field; Mathematics; Data mining; Pure mathematics","score_opus":0.060012365732834784,"score_gpt":0.2942546414822602,"score_spread":0.23424227574942544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2279884163","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011601809,0.019626128,0.38201106,0.007657828,0.0021480974,0.00006699614,0.00044318862,0.0009290974,0.57551575],"genre_scores_gemma":[0.62482655,0.02281847,0.100137144,0.002545859,0.0054334835,0.00029170117,0.0013370326,0.0015361907,0.2410736],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9988777,0.0002607678,0.000080885395,0.0002006939,0.00049279514,0.000087129294],"domain_scores_gemma":[0.99922585,0.00035250143,0.000050289666,0.00018324055,0.00015591535,0.000032259875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011023418,0.0008667599,0.0009122851,0.0018559456,0.0014889047,0.0036178932,0.0010620008,0.0008434682,0.02009763],"category_scores_gemma":[0.0024841959,0.0006182914,0.0009236967,0.002338078,0.0042448468,0.010187673,0.0021310567,0.0042290846,0.006179543],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000032429548,0.0000031422924,0.000015393769,0.000034322125,0.000002786541,0.0000069494154,0.000088737055,0.0001280237,0.00014743958,0.98925215,0.002771895,0.0075458903],"study_design_scores_gemma":[0.0000041574544,0.0000056539866,0.00004952664,0.000025550498,0.00000638768,0.000038625276,0.00002876333,0.0003308253,0.0002372916,0.95320195,0.04606632,0.000004898842],"about_ca_topic_score_codex":0.0006902865,"about_ca_topic_score_gemma":0.0005638255,"teacher_disagreement_score":0.02009763,"about_ca_system_score_codex":0.0024194506,"about_ca_system_score_gemma":0.0008073078,"threshold_uncertainty_score":0.067233324},"labels":[],"label_agreement":null},{"id":"W2280840923","doi":"10.1016/j.net.2015.11.008","title":"A Document-Driven Method for Certifying Scientific Computing Software for Use in Nuclear Safety Analysis","year":2015,"lang":"en","type":"article","venue":"Nuclear Engineering and Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Software engineering; Documentation; Computer science; Traceability; Certification; Programmer; Verification and validation; Programming language; Software requirements specification; Software; Software construction; Software development; Engineering","score_opus":0.031175276997683216,"score_gpt":0.28788694726829556,"score_spread":0.25671167027061237,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2280840923","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011569898,0.000028500133,0.99380624,0.00007571171,0.00002807273,0.00022410335,0.000057386133,0.0036297513,0.0009933171],"genre_scores_gemma":[0.013056792,0.000057577876,0.98235613,0.00005897731,0.000013765354,0.00038566996,0.0003098459,0.0012107396,0.0025505659],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98941994,0.0027342539,0.0009661465,0.0011783555,0.005431667,0.0002696857],"domain_scores_gemma":[0.9778545,0.0072981953,0.001394569,0.0071658525,0.005766981,0.0005198744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008748067,0.0008672857,0.0006558982,0.0034501022,0.0014656535,0.0033293986,0.0029497326,0.0017658721,0.004954359],"category_scores_gemma":[0.02791087,0.0011022793,0.0012927258,0.0016416408,0.001456409,0.0030438951,0.0031327517,0.0029479351,0.003535604],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017957471,0.00042368518,0.0018905047,0.00075052923,0.000075006086,0.0006982018,0.003984596,0.0070607597,0.042639427,0.06545353,0.015718183,0.86112607],"study_design_scores_gemma":[0.00037561532,0.00057069975,0.0024967217,0.0010446463,0.00015890604,0.0040999576,0.0009838311,0.2708915,0.16018091,0.06977154,0.48898625,0.00043931682],"about_ca_topic_score_codex":0.002063687,"about_ca_topic_score_gemma":0.0026417924,"teacher_disagreement_score":0.008748067,"about_ca_system_score_codex":0.0013142701,"about_ca_system_score_gemma":0.0043203393,"threshold_uncertainty_score":0.046264768},"labels":[],"label_agreement":null},{"id":"W2281672363","doi":"10.7287/peerj.preprints.1132v1","title":"Error location in Python: where the mutants hide","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Programming language; Computer science; Python (programming language); Scripting language; Syntax error; Java; Syntax; Abstract syntax; Static analysis; Compiler; Compiled language; Programming paradigm; High-level programming language; Artificial intelligence; Semantics (computer science)","score_opus":0.04883278863831983,"score_gpt":0.30199300757356085,"score_spread":0.253160218935241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2281672363","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12329235,0.0010239473,0.7330556,0.0039041783,0.0011767731,0.00018158079,0.0013697963,0.11749586,0.018499933],"genre_scores_gemma":[0.6561125,0.00080402254,0.26744506,0.002882692,0.00017878169,0.0002527125,0.0013644275,0.053131394,0.017828379],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99560803,0.0008030517,0.00044113523,0.000965169,0.0016867693,0.00049587927],"domain_scores_gemma":[0.9865616,0.00372212,0.0016217562,0.0061888555,0.0013876153,0.0005179175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003304476,0.0009864616,0.00093387946,0.0008112071,0.001450057,0.0027149725,0.0025419814,0.0018226748,0.006786226],"category_scores_gemma":[0.023614531,0.0010463418,0.0012734544,0.0007126809,0.0035148854,0.008713583,0.005356509,0.0036119663,0.003970163],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020308767,0.00069148844,0.041456997,0.0019279844,0.00023775133,0.0054719425,0.0075478028,0.02247271,0.14239673,0.24860615,0.056347378,0.47081214],"study_design_scores_gemma":[0.00015858843,0.00042047873,0.009038985,0.0013372628,0.00030776393,0.004087827,0.0016802404,0.096554026,0.33927828,0.27849376,0.26806593,0.0005768648],"about_ca_topic_score_codex":0.0019639865,"about_ca_topic_score_gemma":0.0021611464,"teacher_disagreement_score":0.006786226,"about_ca_system_score_codex":0.0008781575,"about_ca_system_score_gemma":0.0036396505,"threshold_uncertainty_score":0.022702217},"labels":[],"label_agreement":null},{"id":"W2285418831","doi":"10.7287/peerj.preprints.1771v1","title":"Judging a commit by its cover; or can a commit message predict build failure?","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Commit; Computer science; Code (set theory); Proxy (statistics); Source code; Computer security; Database; Programming language; Machine learning","score_opus":0.013777227243916359,"score_gpt":0.245285266788623,"score_spread":0.23150803954470664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2285418831","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97257894,0.00026731216,0.021351412,0.0009894891,0.000075060874,0.00004311707,0.0010754891,0.00062665873,0.002992531],"genre_scores_gemma":[0.99500704,0.000038608225,0.0036991348,0.000069409965,0.000039786704,0.000014014639,0.00070172275,0.000064641805,0.00036570217],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99590826,0.0010611244,0.00046739882,0.00078182877,0.0013549083,0.00042653747],"domain_scores_gemma":[0.90584624,0.057137102,0.01740227,0.006847434,0.009736785,0.003030185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066422173,0.0006027412,0.0006151563,0.0030763904,0.00054208783,0.0018643804,0.00064974866,0.0011392946,0.0016364409],"category_scores_gemma":[0.087616146,0.0003250594,0.0003220313,0.0021180373,0.0010172059,0.0036862623,0.0015317388,0.0014603455,0.0011222867],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045290845,0.00010619148,0.914763,0.00016470441,0.00016740362,0.000115454044,0.0012250227,0.004501176,0.0052462723,0.00096102303,0.0033950582,0.068901755],"study_design_scores_gemma":[0.00002677884,0.00028969703,0.8790768,0.00010150515,0.00008344834,0.00028173218,0.0021197295,0.10169473,0.006087122,0.006521377,0.0035900103,0.00012697805],"about_ca_topic_score_codex":0.0038495534,"about_ca_topic_score_gemma":0.009296868,"teacher_disagreement_score":0.0066422173,"about_ca_system_score_codex":0.00049681956,"about_ca_system_score_gemma":0.00069835945,"threshold_uncertainty_score":0.03512788},"labels":[],"label_agreement":null},{"id":"W2285477878","doi":"10.1016/j.infsof.2016.01.008","title":"Combining lexical and structural information to reconstruct software layers","year":2016,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Exploit; Software system; Layering; Source code; Software; Software architecture; Process (computing); Software engineering; Documentation; Software maintenance; Programming language","score_opus":0.007871888355526706,"score_gpt":0.23559518383168296,"score_spread":0.22772329547615625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2285477878","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17472032,0.0010182131,0.7987206,0.0005873502,0.00019213115,0.00018096966,0.0037156262,0.008550727,0.012313938],"genre_scores_gemma":[0.55613744,0.0007433719,0.43139362,0.0001234346,0.00008100769,0.00007942752,0.0065856213,0.0014240218,0.003432097],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99940777,0.00011338096,0.00006360211,0.00016594298,0.00016342763,0.0000858497],"domain_scores_gemma":[0.99653035,0.0013953948,0.00025535826,0.0007858923,0.0009231057,0.000109835135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007669371,0.00091602124,0.00053203583,0.005609478,0.0006663561,0.0026455994,0.0010084185,0.0013226337,0.0063007986],"category_scores_gemma":[0.007517957,0.0008314443,0.0010237894,0.0034572526,0.0005643081,0.0059300517,0.0016702366,0.0014735233,0.0034275854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056415336,0.00033134388,0.019315021,0.0007575761,0.0001925489,0.0009953164,0.001542478,0.017838215,0.08296629,0.027950246,0.008412365,0.83913445],"study_design_scores_gemma":[0.00012368248,0.00026489876,0.015535452,0.00040573886,0.0008224565,0.00091290375,0.0020422866,0.7383983,0.071361214,0.13809977,0.03184916,0.00018412719],"about_ca_topic_score_codex":0.0044700922,"about_ca_topic_score_gemma":0.009269497,"teacher_disagreement_score":0.0063007986,"about_ca_system_score_codex":0.0005953554,"about_ca_system_score_gemma":0.0011940316,"threshold_uncertainty_score":0.021078289},"labels":[],"label_agreement":null},{"id":"W2287464799","doi":"10.11575/prism/31123","title":"The Problems of Large-Scale Refactoring: Learning from Eclipse RCP","year":2015,"lang":"en","type":"article","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Eclipse; Computer science; Task (project management); Quality (philosophy); Code (set theory); Process (computing); Scale (ratio); Software engineering; Data science; Risk analysis (engineering); Programming language; Engineering; Business; Software; Systems engineering","score_opus":0.08266976853655637,"score_gpt":0.3257886242390947,"score_spread":0.24311885570253833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2287464799","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.78219426,0.0021388505,0.16208363,0.03352818,0.00024714947,0.00028599412,0.00035760365,0.0013235376,0.017840723],"genre_scores_gemma":[0.8797088,0.0009683443,0.11285387,0.0016230384,0.000115932206,0.00011179811,0.00052023603,0.0004635898,0.0036344454],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9833853,0.008933669,0.0006796749,0.0018700297,0.004360704,0.00077067665],"domain_scores_gemma":[0.7588931,0.19829652,0.011106573,0.014688778,0.014131383,0.0028836143],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026559,0.0010066115,0.00060019275,0.0015843668,0.002453259,0.0042369333,0.0027332574,0.0026346282,0.0013129077],"category_scores_gemma":[0.14063554,0.0007301825,0.00066318206,0.0012978733,0.00445858,0.009096433,0.0052357065,0.005305482,0.0004259136],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007534465,0.0013782196,0.12830554,0.0013705835,0.00017443602,0.0042734444,0.10206656,0.05362786,0.006704533,0.031081235,0.02723082,0.6430333],"study_design_scores_gemma":[0.00042687118,0.0026323975,0.090517715,0.0032497915,0.00029937777,0.0054693306,0.09625672,0.375625,0.031998497,0.231438,0.16123289,0.00085337966],"about_ca_topic_score_codex":0.008346464,"about_ca_topic_score_gemma":0.017344363,"teacher_disagreement_score":0.026559,"about_ca_system_score_codex":0.0037558419,"about_ca_system_score_gemma":0.003845445,"threshold_uncertainty_score":0.14045906},"labels":[],"label_agreement":null},{"id":"W2289443514","doi":"10.1016/j.jss.2016.02.048","title":"Mining trends and patterns of software vulnerabilities","year":2016,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Toronto Metropolitan University","funders":"","keywords":"Secure coding; Computer science; Computer security; Vulnerability (computing); Software; Security bug; Vulnerability management; Buffer overflow; SQL injection; Software security assurance; Software bug; Vulnerability assessment; World Wide Web; Information security; Operating system","score_opus":0.019134659158093432,"score_gpt":0.2533022777641016,"score_spread":0.23416761860600815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2289443514","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9837061,0.0013683446,0.008321479,0.00033050653,0.00003350506,0.00007933253,0.004396438,0.00037721512,0.0013870625],"genre_scores_gemma":[0.9769836,0.00074941537,0.01516965,0.000051703835,0.000041244388,0.00005200738,0.00594991,0.000050178693,0.00095236005],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9982798,0.00014369292,0.00031001875,0.00041300466,0.000686334,0.00016723204],"domain_scores_gemma":[0.9891053,0.004312818,0.0033997246,0.0007942307,0.0018541141,0.0005338178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001209073,0.0004237688,0.0004059186,0.013708343,0.00056489545,0.0011029883,0.0008755865,0.0006707432,0.0008770713],"category_scores_gemma":[0.008457375,0.0002967434,0.0010368056,0.0068277344,0.0003108967,0.0022022275,0.00080167473,0.0006733766,0.0004302455],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020230305,0.00024180242,0.82563424,0.00043641773,0.00034957947,0.00064744527,0.0006455565,0.0023123836,0.006967418,0.0016689157,0.0032084647,0.15768541],"study_design_scores_gemma":[0.000048131824,0.0005779122,0.8602763,0.00039050187,0.0008776281,0.0059119123,0.0034902773,0.08866042,0.013067918,0.011683865,0.014931599,0.000083626415],"about_ca_topic_score_codex":0.00384535,"about_ca_topic_score_gemma":0.010371333,"teacher_disagreement_score":0.013708343,"about_ca_system_score_codex":0.0004653299,"about_ca_system_score_gemma":0.0012953313,"threshold_uncertainty_score":0.0076459646},"labels":[],"label_agreement":null},{"id":"W2289912058","doi":"10.11575/prism/30491","title":"A Survey Paper on Software Architecture Visualization","year":2008,"lang":"en","type":"article","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software visualization; Computer science; Visualization; Software engineering; Software architecture; Resource-oriented architecture; Software architecture description; Software; Architecture; Data science; Software development; Reference architecture; Software construction; Artificial intelligence; Programming language","score_opus":0.058906220834445876,"score_gpt":0.3261674042883758,"score_spread":0.26726118345392996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2289912058","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015231036,0.5769564,0.24587381,0.017387843,0.0033337243,0.00016467617,0.0011459417,0.0024464184,0.13746016],"genre_scores_gemma":[0.10874444,0.69083923,0.15401728,0.0031643857,0.003571841,0.0002125336,0.0025093323,0.0010703492,0.035870526],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982657,0.00051183405,0.00017640096,0.00027860783,0.0006568993,0.0001106772],"domain_scores_gemma":[0.9935673,0.004279247,0.0002698393,0.00042820824,0.0012795408,0.00017590038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018789682,0.0008260604,0.0009631784,0.007959706,0.00095677876,0.0047290223,0.00080378255,0.0014463953,0.01571254],"category_scores_gemma":[0.009144706,0.00070741004,0.00071159203,0.020054858,0.0012064036,0.009824724,0.001392441,0.0013771954,0.0030808689],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046640267,0.00006632656,0.0028639524,0.0034601237,0.000058000096,0.00019130718,0.0023378544,0.0020010655,0.0015874562,0.13744068,0.092754334,0.7571923],"study_design_scores_gemma":[0.000010801255,0.00005059398,0.0026847837,0.0025213405,0.000056967947,0.0010376434,0.0014734637,0.0037082639,0.0007653847,0.07898481,0.90866786,0.000038156628],"about_ca_topic_score_codex":0.0031625738,"about_ca_topic_score_gemma":0.002999813,"teacher_disagreement_score":0.01571254,"about_ca_system_score_codex":0.0013757843,"about_ca_system_score_gemma":0.0017567347,"threshold_uncertainty_score":0.052563667},"labels":[],"label_agreement":null},{"id":"W2293142483","doi":"10.18293/seke2015-042","title":"Using peak analysis for identifying lagged effects between software metrics","year":2015,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Context (archaeology); Time series; Lag; Identification (biology); Software; Series (stratigraphy); Time lag; Data mining; Simple (philosophy); Econometrics; Statistics; Machine learning; Mathematics","score_opus":0.09751997057660587,"score_gpt":0.3208275290474762,"score_spread":0.2233075584708703,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2293142483","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15962486,0.00061362924,0.834865,0.00016468186,0.00015137358,0.00014415296,0.00081273244,0.0018027071,0.0018209778],"genre_scores_gemma":[0.8193314,0.00028644744,0.17827696,0.000060656166,0.0001434209,0.0001459037,0.00081894343,0.00016877819,0.0007674794],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99866307,0.0003684484,0.00011067957,0.00043391358,0.00027788905,0.00014601847],"domain_scores_gemma":[0.98951614,0.0071933027,0.0009546273,0.001061178,0.0010013419,0.00027334623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041596517,0.0009480109,0.0007870057,0.004841893,0.0004994177,0.0014663796,0.0008205156,0.00063647504,0.0025232334],"category_scores_gemma":[0.013018144,0.00040504037,0.000989132,0.0031947948,0.00050844165,0.0022559445,0.0012340801,0.0010214662,0.00054104894],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021964544,0.0009579282,0.18430078,0.0009121266,0.0013514779,0.0014810996,0.0017956856,0.051784236,0.09106804,0.035902653,0.004026952,0.6242226],"study_design_scores_gemma":[0.0001884909,0.0013007745,0.23277652,0.00014466522,0.001019342,0.00094278937,0.000967098,0.6268443,0.043041192,0.08092669,0.011503553,0.00034466668],"about_ca_topic_score_codex":0.0023957882,"about_ca_topic_score_gemma":0.0020034362,"teacher_disagreement_score":0.004841893,"about_ca_system_score_codex":0.00046802528,"about_ca_system_score_gemma":0.0008857732,"threshold_uncertainty_score":0.021998584},"labels":[],"label_agreement":null},{"id":"W2293157443","doi":"10.1145/508386.508389","title":"Explicit programming","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Modular programming; Separation of concerns; Programming language; Source code; Vocabulary; Reusability; Code (set theory); Point (geometry); Software","score_opus":0.03540660668555613,"score_gpt":0.2557462605673754,"score_spread":0.2203396538818193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2293157443","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016072,0.0005486016,0.9565799,0.0011599441,0.0002519651,0.00015130462,0.00025037891,0.0036035948,0.03584704],"genre_scores_gemma":[0.060498558,0.0020560357,0.87678665,0.0016779916,0.00039470615,0.0007551459,0.0010849917,0.0030499962,0.053695846],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99520665,0.00146383,0.0004459385,0.0010662592,0.0014275544,0.00038971542],"domain_scores_gemma":[0.98976445,0.0046637636,0.0005658365,0.003507308,0.0012290165,0.0002696353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005225757,0.001470135,0.0007679514,0.0011584802,0.0016444394,0.005897347,0.0040389025,0.0020102446,0.020740785],"category_scores_gemma":[0.015975837,0.0014300634,0.0019626317,0.0012959451,0.0041964552,0.010847245,0.006606997,0.0050905244,0.0084410645],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005625246,0.00005017326,0.0005069669,0.00047857294,0.00003761504,0.00017275069,0.0010096607,0.0020489725,0.0020442668,0.86921453,0.014746524,0.10963367],"study_design_scores_gemma":[0.00003911786,0.0000443151,0.0001582324,0.00031324296,0.00005989866,0.00053388363,0.00015886444,0.008597838,0.004484062,0.4368342,0.54872674,0.000049735885],"about_ca_topic_score_codex":0.0013157881,"about_ca_topic_score_gemma":0.0016804257,"teacher_disagreement_score":0.020740785,"about_ca_system_score_codex":0.0012293028,"about_ca_system_score_gemma":0.0028195318,"threshold_uncertainty_score":0.06938481},"labels":[],"label_agreement":null},{"id":"W2294018535","doi":"10.1109/icmla.2015.55","title":"What to Learn Next: Recommending Commands in a Feature-Rich Environment","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Feature (linguistics); Task (project management); Baseline (sea); Set (abstract data type); Eclipse; Conjunction (astronomy); Range (aeronautics); Human–computer interaction; Information retrieval; World Wide Web; Machine learning; Programming language","score_opus":0.0505778492301849,"score_gpt":0.28684204067111885,"score_spread":0.23626419144093394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2294018535","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6250799,0.0029194625,0.3425805,0.0020798799,0.00021627214,0.0005100439,0.004120193,0.017470958,0.0050228178],"genre_scores_gemma":[0.6183628,0.0006202396,0.3689439,0.00045642385,0.000103769,0.00012933217,0.006337171,0.00027373564,0.004772578],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99930036,0.00015432318,0.000058183694,0.00028591507,0.00013492141,0.00006633833],"domain_scores_gemma":[0.99671626,0.002102833,0.00022101286,0.00029928598,0.00043026012,0.00023035218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013542668,0.00116961,0.0009206135,0.0018018493,0.0005293447,0.00083627005,0.001587674,0.0013419138,0.0012153098],"category_scores_gemma":[0.0057522072,0.00039663116,0.0006374643,0.0014551503,0.00034574367,0.0023340685,0.00039586963,0.001014703,0.0011248555],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010497681,0.0012957919,0.09023903,0.00052462734,0.0002828315,0.00041353994,0.00052980357,0.054825153,0.007915947,0.0010032235,0.02900678,0.8129135],"study_design_scores_gemma":[0.00018294864,0.0005911453,0.026431296,0.000081231294,0.00017242285,0.0006038944,0.0004441836,0.9449256,0.00989369,0.004461318,0.012131301,0.00008094964],"about_ca_topic_score_codex":0.024182966,"about_ca_topic_score_gemma":0.052727252,"teacher_disagreement_score":0.024182966,"about_ca_system_score_codex":0.00052001467,"about_ca_system_score_gemma":0.0010826504,"threshold_uncertainty_score":0.048084438},"labels":[],"label_agreement":null},{"id":"W2294742457","doi":"10.1007/978-3-319-28406-4_1","title":"How to Build a Recommendation System for Software Engineering","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of British Columbia Hospital","funders":"","keywords":"Computer science; Recommender system; Workflow; Software engineering; Software; Software development; Software system; Software construction; World Wide Web; Database; Programming language","score_opus":0.028953439651432284,"score_gpt":0.2633819535725717,"score_spread":0.23442851392113942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2294742457","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01064135,0.0017469029,0.88776344,0.0021279017,0.0007235603,0.00069750316,0.0049392134,0.06938595,0.021974161],"genre_scores_gemma":[0.034741074,0.0010376374,0.92571783,0.0005363356,0.00015826992,0.0002708223,0.005671503,0.0013045619,0.030561827],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99937963,0.000100831334,0.000059329635,0.00015506409,0.00025843314,0.00004665558],"domain_scores_gemma":[0.99880886,0.00029248724,0.000042297466,0.000353613,0.00042238267,0.00008041757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009578901,0.0007159925,0.0008135166,0.0016884948,0.00076928927,0.0018567682,0.0010889993,0.0016344073,0.021007882],"category_scores_gemma":[0.0046673263,0.0006221433,0.0011817511,0.0016619832,0.00016542817,0.0037952524,0.0008044936,0.0013878339,0.024050724],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010507278,0.0001456828,0.0020567877,0.00026997263,0.00012252743,0.00011336443,0.00009334856,0.0028772803,0.0070322347,0.004238605,0.13824907,0.84469604],"study_design_scores_gemma":[0.00023214842,0.0003130833,0.00960106,0.0004210883,0.0005416356,0.0013616165,0.00037929928,0.36939737,0.040887482,0.049719945,0.52688557,0.00025971734],"about_ca_topic_score_codex":0.0077154413,"about_ca_topic_score_gemma":0.01317521,"teacher_disagreement_score":0.021007882,"about_ca_system_score_codex":0.00042286853,"about_ca_system_score_gemma":0.00060965086,"threshold_uncertainty_score":0.070278406},"labels":[],"label_agreement":null},{"id":"W2295532482","doi":"10.71781/10718","title":"Analysing artefacts dependencies to evolving software systems","year":2013,"lang":"en","type":"book","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Computer science; Software maintenance; Macro; Principal (computer security); Software engineering; Asynchrony (computer programming); Change impact analysis; Software development; Software; Data science; Programming language; Computer security","score_opus":0.046658805490929114,"score_gpt":0.29530401169176185,"score_spread":0.24864520620083275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2295532482","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82849765,0.00086483243,0.15501657,0.0004305694,0.000047098485,0.0002946349,0.0011737685,0.0011316547,0.012543287],"genre_scores_gemma":[0.92874944,0.00041555034,0.06344282,0.000043167583,0.000018563758,0.00019129131,0.0025246625,0.00043897855,0.0041755554],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966684,0.0009840825,0.0002391015,0.00041788482,0.0014331212,0.0002574964],"domain_scores_gemma":[0.97234625,0.018483713,0.0024576413,0.0022550907,0.0040227915,0.00043446646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028740156,0.00050826406,0.0003553656,0.0040665786,0.00062841317,0.0013576677,0.00084804726,0.00069409475,0.0028364463],"category_scores_gemma":[0.035175774,0.0006927731,0.0007784093,0.0026287609,0.000806555,0.0015998672,0.0013388233,0.00096341653,0.0005286343],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007753369,0.000443611,0.2119804,0.0014337867,0.00039968285,0.0047325525,0.012784895,0.20839766,0.050055437,0.043441027,0.005206249,0.46034932],"study_design_scores_gemma":[0.000070572954,0.0003736504,0.23531428,0.00030140762,0.00042443877,0.0013663921,0.003178131,0.655646,0.028662415,0.034530893,0.04001364,0.000118222626],"about_ca_topic_score_codex":0.018142719,"about_ca_topic_score_gemma":0.018006733,"teacher_disagreement_score":0.018142719,"about_ca_system_score_codex":0.0016473895,"about_ca_system_score_gemma":0.0019746197,"threshold_uncertainty_score":0.03607428},"labels":[],"label_agreement":null},{"id":"W2295745251","doi":"","title":"CrashAutomata: an approach for the detection of duplicate crash reports based on generalizable automata","year":2015,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Crash; Computer science; False positive paradox; Software; Automaton; Precision and recall; Data mining; Generalization; Machine learning; Artificial intelligence; Programming language","score_opus":0.030358576573996543,"score_gpt":0.2547864409965724,"score_spread":0.22442786442257584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2295745251","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07726238,0.00068995595,0.88333076,0.00018734542,0.00011491701,0.00043240338,0.0015773742,0.03530473,0.0011001696],"genre_scores_gemma":[0.45037094,0.0003250599,0.5424873,0.00022749469,0.00007882881,0.00044165854,0.0032502396,0.0008700157,0.0019485031],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971621,0.00047010396,0.00034465792,0.0008933895,0.0009312489,0.00019854226],"domain_scores_gemma":[0.9893727,0.0040861857,0.001662312,0.0025745188,0.0020313999,0.00027290473],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017384124,0.001888786,0.0015560663,0.006810285,0.0011936562,0.001930565,0.0026935413,0.0016846596,0.0011006974],"category_scores_gemma":[0.011882585,0.00090482045,0.002009909,0.0024950835,0.00093343476,0.0033425686,0.002288879,0.001735175,0.0012083712],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010891454,0.0006636142,0.12122573,0.0011077501,0.0009106406,0.0018197931,0.0025863044,0.052900854,0.03989715,0.009205269,0.015142854,0.7534509],"study_design_scores_gemma":[0.000050806877,0.00038713458,0.012271187,0.00010414341,0.00031173127,0.0017479395,0.0007877138,0.9115747,0.03801469,0.020794079,0.013790262,0.0001654317],"about_ca_topic_score_codex":0.0070560244,"about_ca_topic_score_gemma":0.010260668,"teacher_disagreement_score":0.0070560244,"about_ca_system_score_codex":0.00090927345,"about_ca_system_score_gemma":0.002128273,"threshold_uncertainty_score":0.01402992},"labels":[],"label_agreement":null},{"id":"W2295844862","doi":"10.1109/asew.2015.28","title":"Analytics for Software Project Management -- Where are We and Where do We Go?","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software project management; Computer science; Status quo; Software analytics; Analytics; Project management; Software; Software development process; Software engineering; Process management; Data science; Software development; Engineering management; Knowledge management; Engineering; Systems engineering; Software construction","score_opus":0.05619973563338226,"score_gpt":0.30879338247440347,"score_spread":0.2525936468410212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2295844862","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23211433,0.17978522,0.3694543,0.14605379,0.0011105074,0.00096080644,0.0026127095,0.0019247208,0.06598367],"genre_scores_gemma":[0.7228432,0.062114645,0.2078928,0.0024444526,0.00053604576,0.0007017822,0.0013700045,0.00027687827,0.0018201454],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96130985,0.027772052,0.0022977903,0.0011916113,0.006702312,0.00072634534],"domain_scores_gemma":[0.7858983,0.17279078,0.011357469,0.01183459,0.015680127,0.0024387087],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.035308298,0.0009574563,0.00087084936,0.014923557,0.002342458,0.012602725,0.0014422333,0.0013072268,0.0019435885],"category_scores_gemma":[0.09691062,0.0004934977,0.0008852033,0.020204987,0.0048300084,0.026716126,0.00417404,0.0022969588,0.00070634735],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012640773,0.00019193778,0.062111333,0.0042395326,0.00016091819,0.00015175785,0.01528069,0.0026915234,0.0015298742,0.15708186,0.012694538,0.74373955],"study_design_scores_gemma":[0.000033505603,0.0004529222,0.055631194,0.018281009,0.00018764626,0.00076640176,0.08536269,0.016824467,0.004208833,0.60046214,0.21754198,0.0002471961],"about_ca_topic_score_codex":0.0023755073,"about_ca_topic_score_gemma":0.002147086,"teacher_disagreement_score":0.9646917,"about_ca_system_score_codex":0.002584684,"about_ca_system_score_gemma":0.0067163156,"threshold_uncertainty_score":0.18673038},"labels":[],"label_agreement":null},{"id":"W2300302512","doi":"10.1109/iwsm-mensura.2011.44","title":"Using Efficient Machine-Learning Models to Assess Two Important Quality Factors: Maintainability and Reusability","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Maintainability; Reusability; Computer science; Software engineering; Software quality; Quality (philosophy); Machine learning; Software metric; Software; Artificial intelligence; Software development; Programming language","score_opus":0.23708793767545858,"score_gpt":0.3699730805243459,"score_spread":0.13288514284888733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2300302512","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07554661,0.00029500492,0.92169815,0.00018333473,0.000015584656,0.00013593626,0.00023939079,0.0007434169,0.0011425392],"genre_scores_gemma":[0.69796354,0.00037448877,0.2993023,0.00005040395,0.000033438577,0.0005651194,0.0007502096,0.00007277271,0.0008876948],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99786633,0.000968927,0.00016717585,0.0003858967,0.00048959884,0.00012207424],"domain_scores_gemma":[0.9818326,0.014850721,0.0012619735,0.0007361672,0.0012290339,0.00008945545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039264862,0.0020723646,0.001238349,0.0029188793,0.00043422176,0.0021375064,0.0013871527,0.0013027751,0.0010499867],"category_scores_gemma":[0.02268287,0.0004919753,0.0009929395,0.0018565597,0.00071086345,0.0025099039,0.0007750896,0.0013387538,0.00049153983],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008755152,0.0002187026,0.008694185,0.00011710929,0.00011708159,0.000033236356,0.00008294208,0.9020925,0.0012237495,0.0040450864,0.00044472088,0.0828432],"study_design_scores_gemma":[0.000006094151,0.000029752786,0.0007020071,0.000011975453,0.00001640229,0.0000109138655,0.000011612345,0.99493897,0.00071364036,0.003435275,0.00011491169,0.000008451697],"about_ca_topic_score_codex":0.0050893915,"about_ca_topic_score_gemma":0.004304582,"teacher_disagreement_score":0.0050893915,"about_ca_system_score_codex":0.001392892,"about_ca_system_score_gemma":0.0011100152,"threshold_uncertainty_score":0.020765483},"labels":[],"label_agreement":null},{"id":"W2314352172","doi":"10.1142/9781860948732_0033","title":"CONSENSUS CONTACT PREDICTION BY LINEAR PROGRAMMING","year":2007,"lang":"en","type":"article","venue":"Computational Systems Bioinformatics","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Linear programming; Algorithm","score_opus":0.016450985169903776,"score_gpt":0.2646623870308433,"score_spread":0.24821140186093954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2314352172","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024967564,0.000058156675,0.9723598,0.00012396395,0.0000136343,0.000037085698,0.000086204054,0.0012606084,0.0010930891],"genre_scores_gemma":[0.5992819,0.00010800381,0.3942505,0.00015564552,0.000056243967,0.00033063008,0.0008831166,0.0003958094,0.004538159],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99874216,0.0003972185,0.00005398393,0.0003205932,0.00032427945,0.0001618554],"domain_scores_gemma":[0.9972247,0.0017573846,0.00026243465,0.00018921756,0.000450423,0.00011579486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014511538,0.00091648527,0.0014969191,0.0010277985,0.0006870578,0.0010544378,0.001577121,0.000851098,0.0031119182],"category_scores_gemma":[0.0049775536,0.0007532735,0.00087763503,0.0010082364,0.00057049724,0.0014039302,0.0010270317,0.001296047,0.00088068476],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000119981654,0.00008632786,0.0013876507,0.000046039462,0.000036509722,0.000074835116,0.000033587854,0.90775144,0.0012233187,0.0037548323,0.0018933737,0.08359208],"study_design_scores_gemma":[0.0000042045945,0.0000063825264,0.000035020166,8.362046e-7,0.0000017056881,0.0000038614085,0.000003538811,0.9975768,0.000269187,0.0020417366,0.00005494109,0.0000017379585],"about_ca_topic_score_codex":0.00548586,"about_ca_topic_score_gemma":0.0049214386,"teacher_disagreement_score":0.00548586,"about_ca_system_score_codex":0.00077701797,"about_ca_system_score_gemma":0.001644462,"threshold_uncertainty_score":0.010907888},"labels":[],"label_agreement":null},{"id":"W2314647138","doi":"10.1115/detc2015-47383","title":"A Novel Application of Gamification for Collecting High-Level Design Information","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Autodesk (Canada)","funders":"","keywords":"Crowdsourcing; Computer science; Context (archaeology); Human–computer interaction; Object (grammar); Game design; Function (biology); Multimedia; Game mechanics; World Wide Web; Artificial intelligence","score_opus":0.1197905922069246,"score_gpt":0.30472936062177325,"score_spread":0.18493876841484863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2314647138","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054606695,0.00027400412,0.92640734,0.00052088196,0.00008592954,0.0011727144,0.0004766512,0.003246659,0.013209223],"genre_scores_gemma":[0.21377553,0.00019106934,0.77997196,0.00017533758,0.000038174923,0.0012372688,0.00050048495,0.00021524214,0.003894938],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9944477,0.0026125684,0.000265098,0.00096284697,0.0014408092,0.00027099976],"domain_scores_gemma":[0.9866831,0.008400349,0.0007449717,0.0022271203,0.0014654457,0.00047900752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044076266,0.0019447235,0.0011389815,0.0044419733,0.00079283596,0.0024035757,0.0022240798,0.0014052525,0.0054752743],"category_scores_gemma":[0.022907754,0.00058872125,0.00091015105,0.0021543107,0.0010936654,0.0031656795,0.004224286,0.0015979869,0.0013879555],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00212996,0.0020645808,0.016302766,0.0016981346,0.00027321998,0.0013558373,0.010507192,0.012493829,0.05115606,0.04807373,0.0072554895,0.8466893],"study_design_scores_gemma":[0.00055247947,0.0025167102,0.026236655,0.001333774,0.00036258966,0.004117016,0.006409859,0.54758686,0.07629876,0.18332541,0.15058129,0.0006786048],"about_ca_topic_score_codex":0.001023959,"about_ca_topic_score_gemma":0.0018971774,"teacher_disagreement_score":0.0054752743,"about_ca_system_score_codex":0.00058963685,"about_ca_system_score_gemma":0.000901387,"threshold_uncertainty_score":0.023310065},"labels":[],"label_agreement":null},{"id":"W2315416862","doi":"10.1145/2856767.2856814","title":"Exploring User Attitudes Towards Different Approaches to Command Recommendation in Feature-Rich Software","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Documentation; Relevance (law); Task (project management); Set (abstract data type); Feature (linguistics); Software documentation; Software; World Wide Web; Data science; Software development; Software development process; Engineering","score_opus":0.2695037850777936,"score_gpt":0.2929308044539849,"score_spread":0.023427019376191294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2315416862","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99834836,0.000021816588,0.00083538593,0.000119321056,0.0000019115935,0.0000137171,0.000007653108,0.000009385516,0.0006424506],"genre_scores_gemma":[0.99792457,0.00004622479,0.0014116433,0.00012559893,0.0000025045524,0.000021279926,0.000017697084,0.0000068029717,0.00044375652],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99407333,0.0042547053,0.0002429317,0.00031925645,0.0008000284,0.00030980186],"domain_scores_gemma":[0.958394,0.03284832,0.0028898623,0.0011373727,0.0030074858,0.0017229698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009580484,0.00031618148,0.00021910774,0.0010101353,0.0015328965,0.0030563052,0.0006690472,0.0013100205,0.0013866791],"category_scores_gemma":[0.032179616,0.000387538,0.0003783741,0.00053893024,0.0018624823,0.0025889087,0.0014149676,0.0014035879,0.00027495413],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041324412,0.0006440673,0.26687723,0.00022836203,0.00006666122,0.00062011246,0.6859898,0.00058690214,0.010251851,0.0011553875,0.00062358624,0.032542706],"study_design_scores_gemma":[0.00007011483,0.0014576114,0.18001746,0.0001937735,0.00008896049,0.00096608134,0.79261327,0.009914853,0.004812586,0.001896106,0.0077257906,0.00024347914],"about_ca_topic_score_codex":0.0050566173,"about_ca_topic_score_gemma":0.007453377,"teacher_disagreement_score":0.009580484,"about_ca_system_score_codex":0.0010669347,"about_ca_system_score_gemma":0.0006346896,"threshold_uncertainty_score":0.050666988},"labels":[],"label_agreement":null},{"id":"W2316465827","doi":"10.7763/ijcte.2013.v5.742","title":"Metrics and Software Quality Evolution: A Case Study on Open Source Software","year":2013,"lang":"en","type":"article","venue":"International Journal of Computer Theory and Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software evolution; Metric (unit); Software quality; Quality (philosophy); Software metric; Software; Open source software; Subject (documents); Quality assurance; Software system; Software engineering; Data mining; Software development; Software construction; World Wide Web; Programming language; Operations management","score_opus":0.026533559643769276,"score_gpt":0.3073577864647694,"score_spread":0.2808242268210001,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2316465827","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9942427,0.00018912877,0.00448898,0.00015491676,0.000004991931,0.00006906201,0.0000718656,0.000025296993,0.0007529738],"genre_scores_gemma":[0.99102414,0.00015707988,0.008220863,0.000021995247,0.000010109514,0.00004419294,0.0001058818,0.000017863991,0.0003978943],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99334514,0.0036916449,0.00036564,0.0004781199,0.001818231,0.0003011819],"domain_scores_gemma":[0.9313344,0.05469367,0.0056481976,0.0026210058,0.0044918894,0.0012108634],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066301967,0.0004238587,0.0004924569,0.00349101,0.0012768236,0.001251905,0.0010644494,0.0015456607,0.0005192989],"category_scores_gemma":[0.027601084,0.0002623206,0.00064655964,0.0052967374,0.0013305328,0.0018735602,0.0012448938,0.0011966252,0.00009023805],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082377647,0.006936371,0.6233359,0.0013868397,0.0003612681,0.020801568,0.047380622,0.049211994,0.020540833,0.015570715,0.002568998,0.21108107],"study_design_scores_gemma":[0.0002752028,0.0049997806,0.7111247,0.00043018282,0.00036466893,0.008823139,0.032511126,0.18838571,0.027091423,0.008719232,0.017016688,0.00025822406],"about_ca_topic_score_codex":0.008306813,"about_ca_topic_score_gemma":0.011674873,"teacher_disagreement_score":0.008306813,"about_ca_system_score_codex":0.001874893,"about_ca_system_score_gemma":0.00094925915,"threshold_uncertainty_score":0.03506428},"labels":[],"label_agreement":null},{"id":"W2316930373","doi":"10.1109/tse.2016.2550458","title":"Developer Micro Interaction Metrics for Software Defect Prediction","year":2016,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Ministry of Education, Science and Technology; Neurosciences Research Foundation","keywords":"Computer science; Eclipse; Leverage (statistics); Software quality assurance; Software quality; Software bug; Software; Software metric; Software engineering; Source code; Software development; Task (project management); Plug-in; Machine learning; Operating system; Systems engineering","score_opus":0.0218172345002419,"score_gpt":0.25171140673187403,"score_spread":0.22989417223163214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2316930373","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54180336,0.0028341431,0.42687723,0.00059440476,0.00013850188,0.0004286973,0.005089963,0.014388456,0.007845254],"genre_scores_gemma":[0.8668675,0.00027654937,0.12833352,0.0000695132,0.000044118453,0.00031360096,0.0028233684,0.0003589007,0.00091294496],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9939766,0.0015015847,0.00041319602,0.0006917517,0.003219349,0.00019746176],"domain_scores_gemma":[0.9585689,0.022108618,0.00883072,0.0038312296,0.0055983565,0.0010622416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030977554,0.0015352081,0.00079647294,0.005848203,0.00035168862,0.0010307211,0.00087886484,0.0006483331,0.0013661777],"category_scores_gemma":[0.033821605,0.00037735724,0.0004354369,0.004054153,0.00027506636,0.0017621457,0.0010930753,0.0009973885,0.0005979152],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048339934,0.0006256983,0.3305445,0.00051244907,0.00030385703,0.00017689088,0.0006222825,0.064542376,0.018905045,0.0034090856,0.010734591,0.56913984],"study_design_scores_gemma":[0.000057597153,0.0011131468,0.16426767,0.0001404713,0.00014771232,0.00035427234,0.00021815689,0.79847324,0.01909723,0.007268052,0.008717351,0.00014508466],"about_ca_topic_score_codex":0.002936218,"about_ca_topic_score_gemma":0.005377657,"teacher_disagreement_score":0.005848203,"about_ca_system_score_codex":0.00068939803,"about_ca_system_score_gemma":0.0008483418,"threshold_uncertainty_score":0.016382694},"labels":[],"label_agreement":null},{"id":"W2327325926","doi":"10.7287/peerj.preprints.1920v3","title":"Analyzing test driven development based on GitHub evidence","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Commit; Agile software development; Computer science; Test-driven development; Software engineering; Java; Process (computing); Software; Software development; Database; Programming language","score_opus":0.04884592265832392,"score_gpt":0.3013920882778682,"score_spread":0.25254616561954424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2327325926","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9728247,0.0024451104,0.01197372,0.0008028942,0.000030459038,0.00033202255,0.004415322,0.00025908338,0.0069166822],"genre_scores_gemma":[0.9791007,0.0009556451,0.010695545,0.00020217286,0.0000362128,0.00038800974,0.007682703,0.00017934991,0.00075955765],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9540083,0.015827417,0.004346623,0.003401757,0.02086429,0.0015516235],"domain_scores_gemma":[0.5113423,0.33389348,0.058702853,0.04670642,0.046914473,0.0024404787],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.023713054,0.00065040606,0.00076169334,0.024400178,0.0008035073,0.0039068027,0.0024964735,0.0012022754,0.0018610902],"category_scores_gemma":[0.26709527,0.000549934,0.0010033724,0.027299372,0.002271371,0.0030986539,0.0033501608,0.0012089123,0.00049928133],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013013989,0.0005639406,0.81315607,0.002809653,0.0010652835,0.0020603554,0.0053548203,0.007759248,0.0040472047,0.009764862,0.005545318,0.14657192],"study_design_scores_gemma":[0.00021359966,0.0010214731,0.9268218,0.0017637897,0.00079162954,0.0019954112,0.0054873135,0.02428982,0.008819428,0.0064734677,0.02217319,0.00014912657],"about_ca_topic_score_codex":0.009155216,"about_ca_topic_score_gemma":0.00953601,"teacher_disagreement_score":0.97628695,"about_ca_system_score_codex":0.0021316765,"about_ca_system_score_gemma":0.0023649784,"threshold_uncertainty_score":0.12540811},"labels":[],"label_agreement":null},{"id":"W2333815456","doi":"10.1115/detc2015-47541","title":"Automatic Extraction of Function Knowledge From Text","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Autodesk (Canada)","funders":"","keywords":"WordNet; Computer science; Word2vec; Natural language processing; Knowledge base; Artificial intelligence; Leverage (statistics); Parsing; Information retrieval; Function (biology); Knowledge extraction","score_opus":0.04113296902143269,"score_gpt":0.3013955782573264,"score_spread":0.26026260923589367,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2333815456","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06724193,0.0039861538,0.8334436,0.0011965694,0.00030938227,0.0011396314,0.049740802,0.024743548,0.018198447],"genre_scores_gemma":[0.1451065,0.0022792395,0.7609467,0.00032774193,0.00021052528,0.00096322136,0.08433298,0.0011996696,0.0046334337],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983511,0.00027342906,0.00025734218,0.00046327693,0.00056289707,0.00009192732],"domain_scores_gemma":[0.992991,0.0039690323,0.0007477339,0.00061292684,0.0015688518,0.00011055334],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010260601,0.0018904319,0.0009797683,0.015831744,0.0010060142,0.0017110622,0.0012694431,0.0010379322,0.0039060959],"category_scores_gemma":[0.007888319,0.00068540673,0.0013246654,0.0066885254,0.000672226,0.004215848,0.0013550996,0.0012129818,0.0036991776],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016161708,0.00024377795,0.007037128,0.0033321655,0.00015636583,0.002189728,0.0015137211,0.0042182356,0.043148596,0.011051684,0.04170612,0.88524085],"study_design_scores_gemma":[0.00015227537,0.0003024654,0.037055675,0.0026240575,0.00073242624,0.0051363055,0.0032298688,0.2194601,0.18458734,0.08679996,0.45954058,0.00037899907],"about_ca_topic_score_codex":0.003580105,"about_ca_topic_score_gemma":0.005272841,"teacher_disagreement_score":0.015831744,"about_ca_system_score_codex":0.0010534177,"about_ca_system_score_gemma":0.002605136,"threshold_uncertainty_score":0.013067186},"labels":[],"label_agreement":null},{"id":"W2334398149","doi":"10.1109/ms.2009.118","title":"Theory of Relative Dependency:&amp;#xD; Higher Coupling Concentration in Smaller Modules and its Implications for Software Refactoring and Quality","year":2009,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Code refactoring; Dependency (UML); Software quality; Software engineering; Quality (philosophy); Software; Computer science; Reliability engineering; Coupling (piping); Programming language; Engineering; Software development; Physics; Mechanical engineering","score_opus":0.08567194059581427,"score_gpt":0.33744661089834255,"score_spread":0.25177467030252826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2334398149","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21752359,0.0034635891,0.64915746,0.0146899875,0.00038414486,0.00027357455,0.0006220106,0.0007369208,0.11314873],"genre_scores_gemma":[0.934377,0.0010520928,0.056294855,0.00212055,0.00030027961,0.00036740783,0.00021680392,0.00018107961,0.005089896],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.994381,0.0018383954,0.00033612436,0.0015684663,0.0015056875,0.00037029857],"domain_scores_gemma":[0.9595183,0.024550062,0.007463569,0.0035158873,0.0034873954,0.0014647741],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006192817,0.0009119201,0.00071674684,0.0037784849,0.0014438756,0.002558784,0.0019570692,0.0020530294,0.012191715],"category_scores_gemma":[0.03534805,0.000827443,0.001331367,0.0030162225,0.007810489,0.0072452137,0.003636346,0.0018983,0.0015519209],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032779158,0.0004959328,0.115744844,0.0007806238,0.00038476085,0.0008218402,0.006658814,0.011343585,0.0050424733,0.732819,0.008103993,0.11747633],"study_design_scores_gemma":[0.00016642899,0.0004317733,0.15136524,0.00030724512,0.00030262774,0.00260084,0.0016100797,0.038648795,0.0030091007,0.77636427,0.025048831,0.00014479224],"about_ca_topic_score_codex":0.0034944522,"about_ca_topic_score_gemma":0.0014303249,"teacher_disagreement_score":0.012191715,"about_ca_system_score_codex":0.0030170442,"about_ca_system_score_gemma":0.0017863724,"threshold_uncertainty_score":0.040785372},"labels":[],"label_agreement":null},{"id":"W2335848699","doi":"","title":"Methods for evaluating, selecting and improving software clustering algorithms","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"York University","funders":"","keywords":"Computer science; Cluster analysis; Data mining; Software metric; Software; Software system; Comparability; Software sizing; Software construction; Software quality; Process (computing); Algorithm; Software development; Construct (python library); Machine learning; Mathematics","score_opus":0.04236197617319268,"score_gpt":0.40600136699169415,"score_spread":0.3636393908185015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2335848699","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06032957,0.003413868,0.9228286,0.0005095738,0.00023542426,0.0008486803,0.0006435955,0.0049094367,0.006281346],"genre_scores_gemma":[0.21414627,0.00084663375,0.7802924,0.00012286894,0.00013053499,0.000591018,0.0015208165,0.00044703373,0.0019024607],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98623395,0.005793055,0.0012248929,0.0020038618,0.004207738,0.0005364479],"domain_scores_gemma":[0.95084655,0.032913763,0.0025923445,0.0037700408,0.009219268,0.00065801863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015326887,0.0027567754,0.002432819,0.009612457,0.0014853333,0.0048067635,0.0035548648,0.0032967879,0.004172784],"category_scores_gemma":[0.07726907,0.0011327324,0.0017858725,0.006296471,0.0012838502,0.005248263,0.0020743,0.0018842688,0.0017009467],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000503966,0.0004146057,0.012534302,0.0005621008,0.0004287121,0.00005263293,0.00019594561,0.24404447,0.0029363954,0.011325284,0.008314504,0.71868706],"study_design_scores_gemma":[0.00009067655,0.00017787804,0.0016316844,0.000075027885,0.00011093395,0.00005043652,0.00011880565,0.98281866,0.0023387528,0.010801183,0.0017542402,0.000031683452],"about_ca_topic_score_codex":0.01244865,"about_ca_topic_score_gemma":0.016927898,"teacher_disagreement_score":0.015326887,"about_ca_system_score_codex":0.0035485616,"about_ca_system_score_gemma":0.0044356957,"threshold_uncertainty_score":0.08105731},"labels":[],"label_agreement":null},{"id":"W2336902201","doi":"10.1007/978-3-319-24282-8","title":"Discovery Science: 18th International Conference, DS 2015, Banff, AB, Canada, October 4-6, 2015. Proceedings","year":2015,"lang":"en","type":"book","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Library science; Event (particle physics); Computer science; Data science; Physics; Astrophysics","score_opus":0.020633226760868102,"score_gpt":0.2801081404951915,"score_spread":0.2594749137343234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2336902201","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050109965,0.31847477,0.14061245,0.03894658,0.060777668,0.00044716842,0.012556767,0.019277243,0.4038964],"genre_scores_gemma":[0.008170999,0.1491824,0.028722141,0.0037517827,0.0041040643,0.00014842977,0.010503387,0.003110708,0.79230607],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99831676,0.00019614096,0.000070035,0.00023074649,0.0010308748,0.00015555645],"domain_scores_gemma":[0.99593025,0.00089414715,0.00010169531,0.00042999708,0.0017424909,0.00090129883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053373724,0.0016863564,0.0016082474,0.004090707,0.0015648826,0.00882674,0.002573682,0.0016656932,0.11475287],"category_scores_gemma":[0.0049241926,0.000986777,0.0008394699,0.0051322924,0.0015841337,0.007329065,0.003663418,0.003049583,0.07103121],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003858506,0.000028423341,0.00009262141,0.0003490747,0.000014553304,0.000032000742,0.00005664377,0.00022217027,0.00077579584,0.0074065127,0.8730436,0.11794005],"study_design_scores_gemma":[0.0000065423833,0.000016366357,0.00020624747,0.00021638989,0.00001017766,0.00012071255,0.00006164453,0.0004474909,0.00051296747,0.004160685,0.9942281,0.000012679234],"about_ca_topic_score_codex":0.01505968,"about_ca_topic_score_gemma":0.050397377,"teacher_disagreement_score":0.9849403,"about_ca_system_score_codex":0.005057428,"about_ca_system_score_gemma":0.007900751,"threshold_uncertainty_score":0.38388658},"labels":[],"label_agreement":null},{"id":"W2339265271","doi":"","title":"Productivity-Based Software Estimation Model: An Economics Perspective and an Empirical Study","year":2014,"lang":"en","type":"article","venue":"International Conference on Software Engineering Advances","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Estimation; Productivity; Computer science; Perspective (graphical); Econometrics; Empirical research; Economics; Macroeconomics; Statistics; Artificial intelligence; Mathematics; Management","score_opus":0.03628275940788956,"score_gpt":0.3461561581016258,"score_spread":0.30987339869373626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2339265271","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1675097,0.0015042844,0.82047194,0.0013258278,0.000060996244,0.00007106594,0.00020955697,0.00023283297,0.008613708],"genre_scores_gemma":[0.9589766,0.0010993895,0.03427756,0.000076612465,0.000098313736,0.00007517193,0.0002409699,0.000045927492,0.005109495],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984889,0.0007541097,0.00006354969,0.00027436807,0.00028539644,0.00013357993],"domain_scores_gemma":[0.98496914,0.01212592,0.0011074615,0.00058364077,0.0010476648,0.00016614576],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045243884,0.0009372954,0.00086773164,0.0020523807,0.0005163276,0.0028768007,0.002248276,0.0019987386,0.0034368737],"category_scores_gemma":[0.020494834,0.00049473287,0.0009774803,0.0025257114,0.0012138654,0.003867116,0.001097089,0.0013214747,0.0006626169],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011055721,0.0002454775,0.019014645,0.00015290848,0.0001396475,0.00030543312,0.0002686307,0.7478194,0.0010469223,0.15728267,0.0018227879,0.07179096],"study_design_scores_gemma":[0.000009805629,0.00004958659,0.0023338217,0.000023728573,0.00004809779,0.000081937964,0.000049388134,0.95700246,0.00033019806,0.03954346,0.00051116914,0.00001634464],"about_ca_topic_score_codex":0.005386782,"about_ca_topic_score_gemma":0.0029638684,"teacher_disagreement_score":0.005386782,"about_ca_system_score_codex":0.0019862708,"about_ca_system_score_gemma":0.001369499,"threshold_uncertainty_score":0.02392757},"labels":[],"label_agreement":null},{"id":"W2339737717","doi":"10.5753/sbsi.2016.5969","title":"Does Technical Debt Lead to the Rejection of Pull Requests?","year":2016,"lang":"en","type":"preprint","venue":"Anais do Simpósio Brasileiro de Sistemas de Informação (SBSI)","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Semtech (Canada)","funders":"Fundação de Amparo à Pesquisa do Estado de Minas Gerais; Conselho Nacional de Desenvolvimento Científico e Tecnológico","keywords":"Technical debt; Convention; Debt; Documentation; Identification (biology); Technical documentation; Computer science; Focus (optics); Business; Risk analysis (engineering); Accounting; Finance; Software; Law; Software development; Political science","score_opus":0.019737057339554512,"score_gpt":0.29330603618907525,"score_spread":0.27356897884952075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2339737717","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9776534,0.00061315723,0.0058828886,0.0014437035,0.00004699491,0.00010697435,0.00018855202,0.00017289394,0.013891491],"genre_scores_gemma":[0.99565375,0.00023411956,0.0014467567,0.0003075604,0.0000590828,0.00005134069,0.00021427278,0.00006617761,0.0019667812],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9780799,0.007683089,0.0028349042,0.001802566,0.007591403,0.0020081473],"domain_scores_gemma":[0.6752589,0.19165885,0.08566026,0.014710406,0.024944628,0.007767057],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.017610567,0.00051515055,0.0005640822,0.003370722,0.0019301405,0.0031486945,0.0011579198,0.0018417264,0.0060851946],"category_scores_gemma":[0.16936095,0.0005050754,0.0006091885,0.0025669106,0.0013516569,0.0045735426,0.0025665553,0.0023201474,0.0012467311],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058646826,0.00034039508,0.88111436,0.00051047455,0.00009712779,0.001742836,0.020823805,0.0005300745,0.005477624,0.0034464246,0.0024480263,0.082882345],"study_design_scores_gemma":[0.000063283864,0.0005817215,0.89836425,0.00064967084,0.00018075807,0.003951575,0.046018783,0.0087694945,0.0058269426,0.009916686,0.025552833,0.00012400841],"about_ca_topic_score_codex":0.0030574896,"about_ca_topic_score_gemma":0.0026690164,"teacher_disagreement_score":0.98238945,"about_ca_system_score_codex":0.0016648617,"about_ca_system_score_gemma":0.0022070347,"threshold_uncertainty_score":0.0931347},"labels":[],"label_agreement":null},{"id":"W2341384651","doi":"10.1145/2911451.2911510","title":"Engineering Quality and Reliability in Technology-Assisted Review","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":96,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Reliability (semiconductor); Reliability engineering; Quality (philosophy); Computer science; Measure (data warehouse); Software quality; Taguchi methods; Data mining; Machine learning; Engineering; Software","score_opus":0.02710917404402359,"score_gpt":0.3098347032712329,"score_spread":0.28272552922720934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2341384651","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04569474,0.00647238,0.9345844,0.003168295,0.00020841176,0.0008700797,0.00025943472,0.0014473101,0.0072949007],"genre_scores_gemma":[0.5582738,0.0021396114,0.43426925,0.00066639954,0.00035286692,0.0010685087,0.00042876787,0.00045287257,0.0023480293],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.7842729,0.15037794,0.013875356,0.009772909,0.039924283,0.0017767069],"domain_scores_gemma":[0.42392996,0.41285643,0.045865025,0.049814925,0.06539743,0.002136271],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.117149316,0.0013811832,0.0021711674,0.007760001,0.0013566834,0.0069790664,0.0027111794,0.0022872582,0.001892586],"category_scores_gemma":[0.48401934,0.0015033745,0.0014562849,0.005476467,0.0031756442,0.007708445,0.0040274444,0.0019445765,0.0011318668],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008916711,0.00026993765,0.028333573,0.0054376987,0.0009037573,0.0004447681,0.0037520851,0.12488944,0.013321931,0.081430845,0.0091229165,0.73120135],"study_design_scores_gemma":[0.00054505793,0.0029330964,0.045127667,0.001998643,0.00097260065,0.0021993231,0.0017824307,0.55888826,0.036384147,0.29614362,0.052378744,0.000646453],"about_ca_topic_score_codex":0.0032853337,"about_ca_topic_score_gemma":0.004005957,"teacher_disagreement_score":0.117149316,"about_ca_system_score_codex":0.0041142157,"about_ca_system_score_gemma":0.0054728836,"threshold_uncertainty_score":0.6195522},"labels":[],"label_agreement":null},{"id":"W2343606968","doi":"10.1007/s10664-015-9416-2","title":"The impact of domain knowledge on the effectiveness of requirements engineering activities","year":2016,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Domain (mathematical analysis); Systems engineering; Computer science; Domain engineering; Engineering; Engineering management; Knowledge management; Software development; Operating system; Mathematics","score_opus":0.02712372993167587,"score_gpt":0.3155331697157048,"score_spread":0.2884094397840289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2343606968","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9867444,0.00049525074,0.0016804041,0.0003822321,0.000012709249,0.00003545371,0.000069132904,0.00004161178,0.010538816],"genre_scores_gemma":[0.99867946,0.00010689718,0.00081161084,0.00003303524,0.0000075168355,0.000011830821,0.000064839565,0.000010815813,0.00027388395],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97673905,0.01551892,0.0013834778,0.0012112479,0.0041555506,0.0009918085],"domain_scores_gemma":[0.25911364,0.70361805,0.01587114,0.009294926,0.009221457,0.0028808904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019523952,0.0004509995,0.00037921246,0.0019668655,0.0006060344,0.0028778473,0.0009775325,0.0014443382,0.0029435726],"category_scores_gemma":[0.29758707,0.00030458948,0.0004795678,0.0010878583,0.00097628805,0.003384607,0.0012740714,0.0015562219,0.00040694088],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009339688,0.009964771,0.4234396,0.0016906499,0.0009880131,0.0008173964,0.003725941,0.06420518,0.017155837,0.008167049,0.0013537619,0.45915216],"study_design_scores_gemma":[0.00070402044,0.007840815,0.84120643,0.0006512025,0.0014529676,0.00086356205,0.0049843933,0.097406544,0.02452833,0.016008131,0.004138869,0.00021477185],"about_ca_topic_score_codex":0.002317965,"about_ca_topic_score_gemma":0.0024367603,"teacher_disagreement_score":0.019523952,"about_ca_system_score_codex":0.0016625143,"about_ca_system_score_gemma":0.0022933455,"threshold_uncertainty_score":0.10325372},"labels":[],"label_agreement":null},{"id":"W2344444819","doi":"10.1145/2902362","title":"On the naturalness of software","year":2016,"lang":"en","type":"article","venue":"Communications of the ACM","topic":"Software Engineering Research","field":"Computer Science","cited_by":291,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Engineering and Physical Sciences Research Council","keywords":"Naturalness; Computer science; Natural language; Natural (archaeology); Software; Artificial intelligence; Program comprehension; Natural language processing; Simplicity; Code (set theory); Programming language; Linguistics; Software system; Set (abstract data type); Epistemology","score_opus":0.04508847840595311,"score_gpt":0.2959249936897846,"score_spread":0.2508365152838315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2344444819","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19995287,0.012385918,0.6270979,0.07718356,0.00073034386,0.00019784392,0.0017350239,0.0015830917,0.07913346],"genre_scores_gemma":[0.87652767,0.0038766246,0.105638266,0.0047338535,0.0008503463,0.0003793408,0.0013885606,0.0006633641,0.0059419214],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9846436,0.007381175,0.0010339842,0.0030011318,0.0035158047,0.000424319],"domain_scores_gemma":[0.9151634,0.06703555,0.0039736484,0.008389581,0.004450383,0.0009875688],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.007878223,0.0007100564,0.00075657084,0.0024999555,0.0027918583,0.0073579787,0.001861116,0.0028797158,0.0035807618],"category_scores_gemma":[0.058342565,0.000902481,0.00080438674,0.0019597765,0.019793943,0.016739609,0.003293452,0.004313029,0.00084826426],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009211133,0.00006342032,0.0054039196,0.000253438,0.00002968558,0.00025453328,0.0043540453,0.0073624314,0.0016077744,0.93702537,0.004947805,0.038605478],"study_design_scores_gemma":[0.000012982461,0.000047798334,0.0028195078,0.0001492934,0.000011924925,0.00042603494,0.0009425919,0.0219327,0.0007132613,0.9463188,0.026584502,0.000040511637],"about_ca_topic_score_codex":0.005746363,"about_ca_topic_score_gemma":0.0035950288,"teacher_disagreement_score":0.9972081,"about_ca_system_score_codex":0.0024365392,"about_ca_system_score_gemma":0.002478671,"threshold_uncertainty_score":0.04166454},"labels":[],"label_agreement":null},{"id":"W2344731017","doi":"10.1016/j.jss.2016.04.058","title":"Missing data techniques in analogy-based software development effort estimation","year":2016,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Analogy; Imputation (statistics); Missing data; Data mining; Computer science; Software; Toleration; Fuzzy logic; Artificial intelligence; Statistics; Machine learning; Mathematics","score_opus":0.03479454045576043,"score_gpt":0.29073897468402865,"score_spread":0.25594443422826824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2344731017","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033962578,0.00024269667,0.9648001,0.00012288564,0.000032132604,0.000046875975,0.000060756392,0.00032859258,0.0004034664],"genre_scores_gemma":[0.67584914,0.00018456334,0.32246783,0.000078603494,0.000056726385,0.00021533054,0.0003058697,0.00009282924,0.00074913114],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9871041,0.008779405,0.00075435604,0.0011930917,0.0018201778,0.0003489045],"domain_scores_gemma":[0.9060247,0.083234034,0.0026990597,0.0044376305,0.00310764,0.00049692916],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014063773,0.00078466855,0.0017775766,0.003330166,0.0008676705,0.001194362,0.003213561,0.0018630552,0.0023958655],"category_scores_gemma":[0.110081255,0.00093413005,0.0012000102,0.0026923262,0.00092891086,0.00399945,0.0029599888,0.002765149,0.00042970284],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00104568,0.0008266545,0.023359539,0.00062308856,0.0005305894,0.00036514044,0.0008724503,0.4234195,0.0033108843,0.04484408,0.0013665697,0.49943593],"study_design_scores_gemma":[0.000041132018,0.00013095573,0.001818566,0.00003577939,0.000048969825,0.00007913997,0.00006872287,0.96705586,0.0010684382,0.029194053,0.00043343406,0.00002500111],"about_ca_topic_score_codex":0.0029491913,"about_ca_topic_score_gemma":0.0030024138,"teacher_disagreement_score":0.9859362,"about_ca_system_score_codex":0.0006548682,"about_ca_system_score_gemma":0.001541443,"threshold_uncertainty_score":0.07437724},"labels":[],"label_agreement":null},{"id":"W2345257595","doi":"10.1007/978-3-319-30282-9_9","title":"Deriving Metrics for Estimating the Effort Needed in Requirements Compliance Work","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Context (archaeology); Requirements engineering; Principal (computer security); Work (physics); Process (computing); Risk analysis (engineering); Key (lock); Function point; Point (geometry); Quality (philosophy); Function (biology); Conformity assessment; Systems engineering; Industrial engineering; Software development; Software; Engineering; Computer security","score_opus":0.06287273047956483,"score_gpt":0.3144839275791002,"score_spread":0.2516111970995354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2345257595","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12803726,0.00074759044,0.8640259,0.00018205505,0.00005714555,0.00043388165,0.00069426786,0.00205611,0.0037658324],"genre_scores_gemma":[0.48345736,0.00030138722,0.5128937,0.000033449214,0.000030340738,0.00047778495,0.0016782326,0.00039238395,0.0007354822],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9728347,0.008841471,0.002878462,0.0017699425,0.012710548,0.00096485973],"domain_scores_gemma":[0.92573744,0.046835296,0.008474688,0.007293038,0.010577919,0.0010815449],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008985212,0.002450551,0.0016464689,0.009887536,0.0005868065,0.003293117,0.002087851,0.0017607395,0.001558878],"category_scores_gemma":[0.09265739,0.001041272,0.0014901366,0.005795258,0.0009518638,0.0045440793,0.0024465234,0.0014939237,0.0007070771],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040291922,0.0005189398,0.04666694,0.0013323199,0.00034403216,0.0003233184,0.000810486,0.22965439,0.017796231,0.024421165,0.0026475713,0.6750818],"study_design_scores_gemma":[0.000036400597,0.0005999581,0.01330146,0.00023563878,0.00012993443,0.00025823154,0.00032607449,0.9563002,0.010267903,0.016653437,0.0018151427,0.00007561472],"about_ca_topic_score_codex":0.0032625038,"about_ca_topic_score_gemma":0.0040715295,"teacher_disagreement_score":0.009887536,"about_ca_system_score_codex":0.002019974,"about_ca_system_score_gemma":0.002136286,"threshold_uncertainty_score":0.04751891},"labels":[],"label_agreement":null},{"id":"W2346502370","doi":"10.1002/smr.1777","title":"An empirical investigation of single‐objective and multiobjective evolutionary algorithms for developer's assignment to bugs","year":2016,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Society of Petroleum Geologists; University of Calgary","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Genetic algorithm; Sorting; Computer science; Operations research; Algorithm; Machine learning; Engineering","score_opus":0.030598698744498356,"score_gpt":0.31648327449455144,"score_spread":0.2858845757500531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2346502370","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95389426,0.00045837203,0.042881336,0.00020466417,0.000022957647,0.00013238304,0.00011619761,0.00012730667,0.0021625299],"genre_scores_gemma":[0.9680562,0.00010383853,0.031181024,0.00002429976,0.0000035076735,0.000088817884,0.00016403926,0.000019631198,0.0003586321],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964579,0.0025326556,0.00013974766,0.00031178183,0.00042581727,0.00013210818],"domain_scores_gemma":[0.9567987,0.03847823,0.001441782,0.0013632511,0.0015595678,0.00035854444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009101906,0.0006012795,0.00056195067,0.0015211594,0.00039316554,0.0013056173,0.0012139839,0.0009467242,0.0013515108],"category_scores_gemma":[0.038375165,0.00026857835,0.00057083316,0.0016976859,0.00043226007,0.0016148476,0.00060205016,0.0008819431,0.000119110555],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005118791,0.0014781931,0.059300672,0.00030130357,0.0002831403,0.00010062028,0.00030762755,0.8026787,0.0010724644,0.0060790745,0.0011808804,0.12670547],"study_design_scores_gemma":[0.000042569794,0.00027623834,0.0057221865,0.00001756041,0.000032770826,0.000040905055,0.00014956627,0.9915121,0.00036144044,0.0014003591,0.0004354715,0.000008829034],"about_ca_topic_score_codex":0.003929324,"about_ca_topic_score_gemma":0.0036368677,"teacher_disagreement_score":0.009101906,"about_ca_system_score_codex":0.0012391779,"about_ca_system_score_gemma":0.0010308636,"threshold_uncertainty_score":0.048136115},"labels":[],"label_agreement":null},{"id":"W2355821122","doi":"","title":"Metrics on Software Maintainability Improvement via Refactoring","year":2009,"lang":"en","type":"article","venue":"Microcomputer applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Code refactoring; Maintainability; Computer science; Software engineering; Software maintenance; Software development; Software; Programmer; Software construction; Software metric; Programming language","score_opus":0.011561134126754994,"score_gpt":0.26594007207195114,"score_spread":0.25437893794519617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2355821122","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69108945,0.008309087,0.28031695,0.00086915947,0.0002030841,0.0005541431,0.00218601,0.0020651673,0.014406981],"genre_scores_gemma":[0.94615734,0.000750785,0.051046122,0.000027337639,0.000037575926,0.00032666663,0.0009551945,0.00006630916,0.00063273223],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98351455,0.006603079,0.001405756,0.000976468,0.007205863,0.00029433332],"domain_scores_gemma":[0.94346493,0.03169413,0.00795908,0.0047217417,0.011485201,0.0006750115],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075669964,0.0007863775,0.0006264962,0.0061592087,0.0003301977,0.0012203052,0.0007385957,0.0006643061,0.0007755439],"category_scores_gemma":[0.062272776,0.00019017508,0.0005284451,0.0062243645,0.0007141294,0.0030905176,0.00084158516,0.0005819347,0.00020573064],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011222531,0.0004620843,0.14911307,0.002527274,0.00053416763,0.00022373085,0.0024485488,0.06998544,0.040449746,0.02696129,0.0030146227,0.7031577],"study_design_scores_gemma":[0.00014963972,0.009251934,0.5769869,0.00088561716,0.00072872784,0.0013003001,0.0013757366,0.26290485,0.08294522,0.035765436,0.027203305,0.0005024363],"about_ca_topic_score_codex":0.0015520361,"about_ca_topic_score_gemma":0.0010490483,"teacher_disagreement_score":0.0075669964,"about_ca_system_score_codex":0.0011820664,"about_ca_system_score_gemma":0.00068159966,"threshold_uncertainty_score":0.040018618},"labels":[],"label_agreement":null},{"id":"W2356036115","doi":"10.1016/j.jss.2016.05.015","title":"Topic-based software defect explanation","year":2016,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Queen's University","funders":"","keywords":"Computer science; Cohesion (chemistry); Software quality; Software metric; Software; Source lines of code; Eclipse; Compiler; Software system; Software engineering; Product metric; Software development; Software analytics; Source code; Data mining; Data science; Software construction; Programming language; Metric space","score_opus":0.017429593696882448,"score_gpt":0.24499590244833766,"score_spread":0.22756630875145523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2356036115","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28336757,0.0024186035,0.67358714,0.0034407282,0.00047225942,0.00081559044,0.008440999,0.011863446,0.015593639],"genre_scores_gemma":[0.8757698,0.0004282159,0.11350263,0.00014907177,0.00012223874,0.00021674599,0.006235168,0.00027048867,0.0033056354],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99813914,0.00060298236,0.00016039479,0.00038068957,0.00055012736,0.0001667155],"domain_scores_gemma":[0.98388255,0.010531827,0.0010416431,0.0010973704,0.003175843,0.00027079417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021368554,0.00070524326,0.0005713399,0.0049799206,0.00065095205,0.001553267,0.001551881,0.0015228993,0.006330301],"category_scores_gemma":[0.02153315,0.00025297204,0.0010431673,0.002861296,0.00036892906,0.0034243471,0.0015944844,0.0011631099,0.00096247613],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016807411,0.00068552844,0.12884608,0.0013344993,0.0005749921,0.00077296945,0.0033962522,0.044724137,0.01064589,0.035555203,0.042908628,0.7288751],"study_design_scores_gemma":[0.00023709422,0.00027924593,0.052387122,0.0002711877,0.0006356332,0.0005570649,0.0019087346,0.8502182,0.009701202,0.06013772,0.023552168,0.00011463333],"about_ca_topic_score_codex":0.009188721,"about_ca_topic_score_gemma":0.0089852335,"teacher_disagreement_score":0.009188721,"about_ca_system_score_codex":0.0010740171,"about_ca_system_score_gemma":0.0017490283,"threshold_uncertainty_score":0.021176994},"labels":[],"label_agreement":null},{"id":"W2360967250","doi":"10.1145/2884781.2884804","title":"Automatically learning semantic features for defect prediction","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":692,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software bug; Focus (optics); ENCODE; Machine learning; Code (set theory); Artificial intelligence; Software; Software engineering; Data mining; Programming language","score_opus":0.010519247014547421,"score_gpt":0.2535673509846479,"score_spread":0.24304810397010046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2360967250","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31424978,0.0016476443,0.65280247,0.0007202987,0.00016870059,0.00022171965,0.0075309016,0.0191849,0.003473699],"genre_scores_gemma":[0.8248921,0.00035506897,0.16045958,0.00011642685,0.00007470138,0.0002114607,0.012749885,0.00034000815,0.00080074195],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99909234,0.00013833791,0.000068741225,0.00029451278,0.00030362318,0.00010247709],"domain_scores_gemma":[0.9963439,0.0019114644,0.00045940117,0.00045667347,0.0007136806,0.000114802984],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006788636,0.0013409959,0.0009724786,0.0039435597,0.00048029033,0.0006027179,0.001027311,0.0012016612,0.0012047384],"category_scores_gemma":[0.0061564366,0.000385203,0.0008928801,0.0020786684,0.0005241603,0.0027870676,0.000903044,0.0011334557,0.0007210067],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070767035,0.0010716195,0.06411071,0.00061740325,0.0001736561,0.00076956465,0.00026168328,0.10327406,0.026624022,0.011264866,0.03357187,0.75755286],"study_design_scores_gemma":[0.00004269535,0.00013970684,0.007829441,0.000051065883,0.00007390664,0.00021978762,0.00008586019,0.9543732,0.009957517,0.022884281,0.0043078815,0.000034695353],"about_ca_topic_score_codex":0.0036973453,"about_ca_topic_score_gemma":0.0058964435,"teacher_disagreement_score":0.0039435597,"about_ca_system_score_codex":0.00066585274,"about_ca_system_score_gemma":0.0010551588,"threshold_uncertainty_score":0.007351637},"labels":[],"label_agreement":null},{"id":"W2362712727","doi":"10.5220/0005822701320139","title":"Systematic Mapping Study of Ensemble Effort Estimation","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Estimation; Systems engineering; Engineering","score_opus":0.02250864086124901,"score_gpt":0.26929628285952484,"score_spread":0.24678764199827583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2362712727","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5015846,0.3530772,0.1189625,0.0026545234,0.0006762159,0.002858782,0.0068900096,0.00041066392,0.012885546],"genre_scores_gemma":[0.9037779,0.051319666,0.03937412,0.00030365997,0.00014292197,0.0016361547,0.0025013085,0.000064348584,0.0008798445],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9610852,0.016613182,0.008998611,0.004305179,0.008414824,0.0005830966],"domain_scores_gemma":[0.65794027,0.23692049,0.03663952,0.018486608,0.049003527,0.0010096125],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04676703,0.0009908751,0.0020265693,0.045084435,0.0007794882,0.003040586,0.0010810462,0.00070936984,0.0015773718],"category_scores_gemma":[0.17551741,0.0005745134,0.0033143533,0.038771614,0.0008420039,0.004836556,0.0025377728,0.0007606173,0.0002805392],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004750222,0.00013617781,0.33043587,0.053803183,0.014480536,0.0009793391,0.0076638972,0.00367779,0.0026216272,0.011324211,0.0028816364,0.5715207],"study_design_scores_gemma":[0.00032452223,0.0028176932,0.6144199,0.10863868,0.056128554,0.004400694,0.034349583,0.028903782,0.018510059,0.0325913,0.0984828,0.0004324952],"about_ca_topic_score_codex":0.0024789423,"about_ca_topic_score_gemma":0.0030898093,"teacher_disagreement_score":0.95323294,"about_ca_system_score_codex":0.0015122225,"about_ca_system_score_gemma":0.006499089,"threshold_uncertainty_score":0.2473306},"labels":[],"label_agreement":null},{"id":"W2364021410","doi":"10.5220/0005857601720180","title":"Source and Test Code Size Prediction - A Comparison between Use Case Metrics and Objective Class Points","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Computer science; Metric (unit); Regression testing; Source lines of code; Source code; Java; Software metric; Data mining; Linear regression; Test case; Parametric statistics; Regression analysis; Code (set theory); Software; Software quality; Machine learning; Software development; Statistics; Set (abstract data type); Programming language; Mathematics; Software construction","score_opus":0.034002242440454696,"score_gpt":0.28004287180510995,"score_spread":0.24604062936465526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2364021410","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96389705,0.00033489108,0.03111898,0.000118471595,0.000017598619,0.00010324451,0.0012650511,0.0006735412,0.002471119],"genre_scores_gemma":[0.9875246,0.00007607513,0.010104764,0.000016605427,0.000012607588,0.00011040632,0.0015227734,0.0000697334,0.0005624389],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99092287,0.003159961,0.00061998266,0.0013323994,0.0037176067,0.0002472755],"domain_scores_gemma":[0.84720665,0.11393264,0.018632617,0.007551066,0.011502471,0.0011745114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0083412705,0.0008260486,0.00051222346,0.005234775,0.00020251564,0.0009951688,0.0009690749,0.000812367,0.0011088622],"category_scores_gemma":[0.07947439,0.00032704975,0.000656736,0.003605703,0.00038823445,0.0024863167,0.00085214956,0.0007951887,0.00066702324],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037154419,0.00048823055,0.83227205,0.00019104456,0.0002346605,0.00009304856,0.0006204673,0.021609439,0.002293743,0.0005962322,0.0010374296,0.14019215],"study_design_scores_gemma":[0.000026872085,0.0007023045,0.7598414,0.00006750812,0.000068655754,0.00021592093,0.00032909334,0.23182002,0.004470652,0.0010146237,0.0013900471,0.000052988227],"about_ca_topic_score_codex":0.0024532918,"about_ca_topic_score_gemma":0.0034246289,"teacher_disagreement_score":0.0083412705,"about_ca_system_score_codex":0.0006363222,"about_ca_system_score_gemma":0.00042227274,"threshold_uncertainty_score":0.044113338},"labels":[],"label_agreement":null},{"id":"W2366520322","doi":"","title":"Feature location based on impact analysis","year":2007,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Feature (linguistics); Component (thermodynamics); TRACE (psycholinguistics); Source code; Feature model; Reverse engineering; Data mining; Dependency (UML); Software; Artificial intelligence; Programming language","score_opus":0.02185341761854571,"score_gpt":0.31011392259258036,"score_spread":0.28826050497403466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2366520322","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06348233,0.00059823645,0.9223013,0.00017972974,0.000091754526,0.00039247898,0.0007985513,0.0060519623,0.0061036036],"genre_scores_gemma":[0.70805436,0.00038101088,0.28642738,0.0000727871,0.00010440947,0.00027543012,0.0013611935,0.00043044664,0.0028930162],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957123,0.00069559884,0.00027888495,0.00067004206,0.0022540414,0.00038908975],"domain_scores_gemma":[0.9877028,0.005856154,0.0017705442,0.0015325355,0.0028805418,0.00025747498],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020005342,0.0020212105,0.001579855,0.016285297,0.00092259847,0.002050282,0.0016824396,0.0013026189,0.005151518],"category_scores_gemma":[0.014849042,0.00045211284,0.0018945147,0.0064176773,0.0010173396,0.0030981044,0.0019458671,0.0012273068,0.0012867294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010230978,0.0005095995,0.06530024,0.0008438944,0.00056986674,0.0024690318,0.0006964802,0.17579363,0.048992462,0.040068615,0.006130256,0.65760285],"study_design_scores_gemma":[0.00005410201,0.0003176263,0.016824452,0.000102733975,0.00031738981,0.0011896705,0.00029409604,0.9028779,0.02981863,0.04094579,0.00707728,0.00018031475],"about_ca_topic_score_codex":0.004938318,"about_ca_topic_score_gemma":0.0033083495,"teacher_disagreement_score":0.016285297,"about_ca_system_score_codex":0.0012955274,"about_ca_system_score_gemma":0.0010267484,"threshold_uncertainty_score":0.017233491},"labels":[],"label_agreement":null},{"id":"W2366532918","doi":"10.1145/2884781.2884800","title":"Augmenting API documentation with insights from stack overflow","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":241,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Documentation; Computer science; Automatic summarization; Software documentation; Set (abstract data type); Software; Artificial intelligence; Natural language processing; Information retrieval; Programming language; Software development; Software development process","score_opus":0.01075786429482343,"score_gpt":0.24120317513156986,"score_spread":0.23044531083674644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2366532918","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5987288,0.0017866845,0.3668258,0.00097602897,0.00018011716,0.00063459313,0.010828587,0.01544913,0.004590264],"genre_scores_gemma":[0.475948,0.0007367904,0.49903667,0.0002314464,0.0001539582,0.00038428025,0.02074355,0.0007825259,0.001982696],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99783057,0.00068631564,0.00034535222,0.00041189426,0.00063260284,0.00009332076],"domain_scores_gemma":[0.97193354,0.016735937,0.003949444,0.0019723233,0.005082923,0.0003258484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002487831,0.0010842867,0.00064410164,0.005494804,0.0004903645,0.0010874843,0.0005835582,0.000673144,0.001070846],"category_scores_gemma":[0.024755908,0.0004670219,0.00060982694,0.002476837,0.00029475248,0.0022498874,0.0015078993,0.0009075172,0.000805862],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054362166,0.00034733472,0.060401835,0.002659591,0.00022897465,0.0019035211,0.0064725457,0.0057081706,0.11907751,0.0014182631,0.0142746605,0.786964],"study_design_scores_gemma":[0.00016141841,0.001429792,0.2692448,0.001140581,0.00096349383,0.0053739143,0.005791989,0.29428396,0.30168614,0.011646236,0.10785513,0.0004225522],"about_ca_topic_score_codex":0.0014263054,"about_ca_topic_score_gemma":0.003541001,"teacher_disagreement_score":0.005494804,"about_ca_system_score_codex":0.00036632564,"about_ca_system_score_gemma":0.0011057947,"threshold_uncertainty_score":0.013157129},"labels":[],"label_agreement":null},{"id":"W2367798545","doi":"10.1145/2884781.2884857","title":"Automated parameter optimization of classification techniques for defect prediction models","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":344,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Classifier (UML); Machine learning; Artificial intelligence; Random forest; Predictive modelling; Software bug; Decision tree; Software; Data mining","score_opus":0.03975154568969916,"score_gpt":0.2891461855171243,"score_spread":0.24939463982742513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2367798545","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.123099804,0.0011951424,0.8607571,0.0007108875,0.00013919716,0.0005794632,0.0006138244,0.01027538,0.0026291492],"genre_scores_gemma":[0.620587,0.0004477721,0.37377843,0.00034115242,0.000081912694,0.000867084,0.0017294057,0.0008540303,0.0013131198],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995197,0.0022023981,0.00044773973,0.0010400887,0.00078104023,0.00033179967],"domain_scores_gemma":[0.97701824,0.014798766,0.002094802,0.0034823564,0.0024302988,0.00017559594],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072777555,0.002870057,0.0015748289,0.0031754258,0.00075319025,0.00196185,0.0021041187,0.0019754404,0.0018476616],"category_scores_gemma":[0.04490184,0.0010023095,0.00193185,0.0019990688,0.00091026287,0.0032599058,0.0013863966,0.0035064798,0.0014828255],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030954447,0.00047762218,0.01534832,0.00025496027,0.00041305413,0.00013682479,0.0002749333,0.5643507,0.008706163,0.0023296855,0.005803585,0.40159455],"study_design_scores_gemma":[0.000036297744,0.00013211156,0.0023597972,0.00004737984,0.000058253252,0.000105342275,0.000074888514,0.98427296,0.0065136454,0.0046910034,0.0016662363,0.00004215843],"about_ca_topic_score_codex":0.0050244955,"about_ca_topic_score_gemma":0.006628972,"teacher_disagreement_score":0.0072777555,"about_ca_system_score_codex":0.0015156433,"about_ca_system_score_gemma":0.0018197976,"threshold_uncertainty_score":0.038488925},"labels":[],"label_agreement":null},{"id":"W2373401668","doi":"10.1145/2884781.2884864","title":"Understanding asynchronous interactions in full-stack JavaScript","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Intel Corporation","keywords":"Callback; JavaScript; Computer science; Asynchronous communication; Unobtrusive JavaScript; Program comprehension; Programming language; Call stack; Server-side; Web application; Context (archaeology); Operating system; Stack (abstract data type); Software engineering; Rich Internet application; Software","score_opus":0.1042894808353669,"score_gpt":0.30302604816573453,"score_spread":0.19873656733036763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2373401668","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19992319,0.000111984526,0.7921764,0.00021378217,0.000019573772,0.000102371414,0.00014330745,0.0054500597,0.0018592936],"genre_scores_gemma":[0.6321539,0.00020583508,0.36444664,0.00010268435,0.000016305326,0.00011515918,0.00033768092,0.0010278837,0.0015938607],"study_design_codex":"bench_or_experimental","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987418,0.000494035,0.00007336507,0.00028219915,0.0003451695,0.00006349108],"domain_scores_gemma":[0.9899928,0.00745841,0.0009055105,0.0007568616,0.00070246967,0.00018394212],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014903325,0.00082220143,0.00042609006,0.00062087836,0.0004580458,0.0018092667,0.0012028557,0.0010214492,0.001562885],"category_scores_gemma":[0.0121344505,0.000739048,0.00054148235,0.00033412818,0.000989222,0.0044232663,0.0013888221,0.0011826645,0.00036877644],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010172485,0.00097506493,0.026413245,0.0012769137,0.000117447664,0.0020077955,0.03231403,0.15902159,0.36682665,0.04859817,0.005487137,0.35594472],"study_design_scores_gemma":[0.000072945935,0.00022314813,0.0084425835,0.00010069909,0.00005397786,0.0004734108,0.0010572413,0.9003242,0.046126835,0.032061137,0.010988138,0.00007565309],"about_ca_topic_score_codex":0.004263661,"about_ca_topic_score_gemma":0.0048125377,"teacher_disagreement_score":0.004263661,"about_ca_system_score_codex":0.0006064475,"about_ca_system_score_gemma":0.0009331937,"threshold_uncertainty_score":0.008477688},"labels":[],"label_agreement":null},{"id":"W2373727499","doi":"","title":"Software Development Model Based on Bayesian Network","year":2007,"lang":"en","type":"article","venue":"Microcomputer applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Bayesian network; Software; Variable-order Bayesian network; Process (computing); Construct (python library); Dynamic Bayesian network; Data mining; Key (lock); Artificial intelligence; Software development; Bayesian probability; Machine learning; Software engineering; Bayesian inference; Programming language","score_opus":0.011717868054659,"score_gpt":0.2527484383912139,"score_spread":0.24103057033655492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2373727499","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0061542834,0.00043007298,0.9862705,0.00038320836,0.000025020883,0.00005782621,0.0001465505,0.0002961938,0.006236506],"genre_scores_gemma":[0.50408405,0.0035687785,0.47386667,0.00019012285,0.00013624252,0.0009202938,0.00089454354,0.0001806794,0.016158609],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9974605,0.0008369484,0.00014242303,0.0004887821,0.00093485875,0.00013650267],"domain_scores_gemma":[0.99769706,0.0013806361,0.00025193623,0.0001011783,0.0004924465,0.00007678104],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021970335,0.00076362735,0.0008151818,0.0024309545,0.00074917224,0.0019385768,0.001977455,0.0013964581,0.0032751502],"category_scores_gemma":[0.006906624,0.00071208394,0.0010516555,0.0019905725,0.00091566215,0.0035240815,0.0009837323,0.0011471388,0.0008363226],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007077968,0.000056943678,0.0022287592,0.00017584977,0.000087299944,0.00032977067,0.00037133074,0.6735777,0.0011879134,0.24601641,0.0025047718,0.07339244],"study_design_scores_gemma":[0.000023218887,0.000022869574,0.00031365253,0.00002758855,0.00003569771,0.00010234449,0.000020031668,0.93245566,0.00028224147,0.06220555,0.0044891387,0.000022051463],"about_ca_topic_score_codex":0.014144898,"about_ca_topic_score_gemma":0.008510309,"teacher_disagreement_score":0.014144898,"about_ca_system_score_codex":0.0018719176,"about_ca_system_score_gemma":0.0018383467,"threshold_uncertainty_score":0.028125107},"labels":[],"label_agreement":null},{"id":"W23742278","doi":"10.1007/978-1-4020-6264-3_29","title":"A Multi-Agent Framework for Building an Automatic Operational Profile","year":2007,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Systems engineering; Engineering","score_opus":0.07510658261435339,"score_gpt":0.3511159639488267,"score_spread":0.27600938133447334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W23742278","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005226521,0.000076449476,0.9938606,0.000057470108,0.00003133303,0.0000437658,0.00007220352,0.0015206318,0.00381491],"genre_scores_gemma":[0.025695326,0.00016548393,0.96656466,0.00004817129,0.000025301104,0.00011321836,0.00026482187,0.00033736287,0.0067856107],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99953127,0.00010492442,0.00004054841,0.00009124855,0.00018967838,0.00004235907],"domain_scores_gemma":[0.9996026,0.0001441785,0.000031690906,0.00009668523,0.00009379235,0.00003100944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009403505,0.0009145226,0.0007212613,0.0009601182,0.0010954838,0.002765546,0.001988304,0.0012789194,0.008395365],"category_scores_gemma":[0.0018259047,0.00079176726,0.0010094119,0.00081289664,0.0006417239,0.0026576903,0.0014828368,0.0020278688,0.004202431],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009230936,0.00010582907,0.00053961755,0.00027360977,0.00008321508,0.0004102734,0.0005573134,0.12725428,0.008614357,0.4420617,0.022379853,0.39762765],"study_design_scores_gemma":[0.000018252606,0.000026307676,0.00015933975,0.000078626545,0.000045490193,0.00021319963,0.00007401709,0.76417804,0.005984594,0.16314885,0.066034,0.000039281473],"about_ca_topic_score_codex":0.0035377264,"about_ca_topic_score_gemma":0.005733211,"teacher_disagreement_score":0.008395365,"about_ca_system_score_codex":0.00077475695,"about_ca_system_score_gemma":0.0011764998,"threshold_uncertainty_score":0.028085351},"labels":[],"label_agreement":null},{"id":"W2374812233","doi":"10.1145/2884781.2884852","title":"Revisiting code ownership and its relationship with software quality in the scope of modern code review","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":136,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; McGill University","funders":"","keywords":"Code review; Computer science; Heuristics; Code (set theory); Software quality; Scope (computer science); KPI-driven code analysis; Object code; Redundant code; Software; Code smell; Quality (philosophy); Software engineering; Programming language; Software development; Code generation; Computer security; Key (lock); Operating system","score_opus":0.09319115713258627,"score_gpt":0.3394747782210961,"score_spread":0.24628362108850985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2374812233","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9852411,0.002052508,0.0071568484,0.00074085285,0.00003153695,0.00008203602,0.00006834112,0.00006699462,0.004559722],"genre_scores_gemma":[0.99814975,0.00017362226,0.0012688021,0.000046991456,0.000025952015,0.000016425585,0.00003892041,0.000017682149,0.0002618355],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.95208764,0.020567317,0.0049089114,0.0052780574,0.01523887,0.001919267],"domain_scores_gemma":[0.18085158,0.5687206,0.18216221,0.023307689,0.03876571,0.006192265],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03922754,0.00041613256,0.0005516704,0.008113853,0.0012871375,0.004402994,0.0013737495,0.000989141,0.0015749],"category_scores_gemma":[0.38753113,0.00053392205,0.00051983277,0.0049920226,0.004301838,0.006663828,0.0035643426,0.0014021945,0.00024221967],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000092044706,0.00005998449,0.9596271,0.00017224741,0.00009847466,0.00016371666,0.0075006867,0.0005853257,0.0006738583,0.00089134503,0.00027439458,0.029860886],"study_design_scores_gemma":[0.000012761487,0.00019311075,0.9838306,0.00019022803,0.00010037486,0.000519454,0.0064354176,0.0031203362,0.001009702,0.0020085122,0.002535988,0.000043569282],"about_ca_topic_score_codex":0.0067859236,"about_ca_topic_score_gemma":0.009690536,"teacher_disagreement_score":0.96077245,"about_ca_system_score_codex":0.0023168053,"about_ca_system_score_gemma":0.0030678718,"threshold_uncertainty_score":0.20745754},"labels":[],"label_agreement":null},{"id":"W2383417445","doi":"10.1145/2884781.2884844","title":"RETracer","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Crash; Computer science; TRACE (psycholinguistics); Software; Semantics (computer science); Computer security; Software bug; World Wide Web; Data science; Programming language","score_opus":0.015953877596840166,"score_gpt":0.25557177350940385,"score_spread":0.2396178959125637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2383417445","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049123936,0.0013399033,0.22653055,0.0027097058,0.002592909,0.00040508466,0.011416113,0.64107406,0.10901923],"genre_scores_gemma":[0.090344414,0.0021799824,0.18128198,0.008374305,0.0010645329,0.0010202844,0.08199626,0.34263054,0.29110768],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9955597,0.00067490217,0.00030541758,0.0012491264,0.0015004764,0.0007104525],"domain_scores_gemma":[0.9934546,0.0012043748,0.00020464495,0.0033286838,0.001509522,0.00029825483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030960653,0.0030814062,0.0017031326,0.0017667506,0.0019820358,0.004909674,0.0070767067,0.0035355503,0.1784323],"category_scores_gemma":[0.017127667,0.0020207067,0.003224272,0.0014929118,0.0015258909,0.011427251,0.0118863415,0.005315542,0.2179055],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007108857,0.00022193283,0.002579076,0.00080854294,0.000111465175,0.00053547666,0.0006460896,0.002634399,0.0036978472,0.039802905,0.78035146,0.16789992],"study_design_scores_gemma":[0.00008553018,0.00006644538,0.00050165463,0.00019580965,0.00005944458,0.0004563796,0.00025166833,0.015354487,0.012145982,0.032244176,0.9385345,0.000103962586],"about_ca_topic_score_codex":0.0062553408,"about_ca_topic_score_gemma":0.0068709585,"teacher_disagreement_score":0.1784323,"about_ca_system_score_codex":0.0015978614,"about_ca_system_score_gemma":0.0030454916,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2385483600","doi":"10.1145/2884781.2884839","title":"Cross-project defect prediction using a connectivity-based unsupervised classifier","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":249,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Homogeneity (statistics); Classifier (UML); Computer science; Reuse; Machine learning; Artificial intelligence; Metric (unit); Data mining; Engineering","score_opus":0.05826272584185882,"score_gpt":0.31995871329536246,"score_spread":0.26169598745350364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2385483600","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57917154,0.0003720894,0.4141143,0.0002311882,0.00008052947,0.00018625232,0.0011040649,0.0020331421,0.0027069151],"genre_scores_gemma":[0.9423076,0.00013150992,0.053728223,0.00003885339,0.00006974049,0.00013204882,0.0021932393,0.00006022766,0.0013385484],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988803,0.0001527904,0.000087533546,0.0003836914,0.00035254747,0.00014305268],"domain_scores_gemma":[0.9957124,0.001522198,0.0006901973,0.000497346,0.0013304586,0.00024737974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011321832,0.0008024444,0.0010895359,0.0043055937,0.0004873714,0.0008806113,0.0014537728,0.0012068931,0.00062160933],"category_scores_gemma":[0.0051679127,0.00025036858,0.0007522621,0.0022612798,0.00032548548,0.0013543649,0.0009652007,0.00075817853,0.00060055096],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005302012,0.0011754421,0.18694043,0.000149579,0.0003702483,0.0006653686,0.0002617039,0.23595594,0.010720463,0.0018304994,0.008243275,0.553157],"study_design_scores_gemma":[0.000011012596,0.000102647646,0.015230485,0.000011635091,0.00004153471,0.0001453217,0.000037852533,0.9815524,0.001303992,0.0010214654,0.00052697456,0.00001471818],"about_ca_topic_score_codex":0.0047922055,"about_ca_topic_score_gemma":0.006382905,"teacher_disagreement_score":0.0047922055,"about_ca_system_score_codex":0.00048793614,"about_ca_system_score_gemma":0.0007975579,"threshold_uncertainty_score":0.009528637},"labels":[],"label_agreement":null},{"id":"W2394837668","doi":"10.1109/saner.2016.103","title":"Do Code Smells Impact the Effort of Different Maintenance Programming Activities?","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code smell; Computer science; Software maintenance; Task (project management); Code (set theory); Java; Empirical research; Software engineering; Software; Programming language; Software quality; Software development; Engineering","score_opus":0.016507489413852217,"score_gpt":0.28431256144801964,"score_spread":0.2678050720341674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2394837668","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9939168,0.00068266055,0.0027738921,0.00033154004,0.000022923805,0.000026068396,0.00062569993,0.00027695633,0.0013435513],"genre_scores_gemma":[0.9965491,0.00014276007,0.0014811871,0.00008095096,0.000020256492,0.000017715183,0.0011630096,0.000117848016,0.00042705605],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9933477,0.0015488784,0.0006954688,0.0015733268,0.0022688012,0.0005657355],"domain_scores_gemma":[0.8307244,0.109286,0.03910443,0.009320771,0.007860281,0.0037040357],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046264958,0.0006117375,0.00048078437,0.0023292045,0.00029327747,0.0017075895,0.0006332437,0.000879947,0.0012956705],"category_scores_gemma":[0.066935144,0.00047330474,0.0007210341,0.0018166335,0.0007603898,0.002716433,0.0009375414,0.0010976493,0.0005887552],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003556858,0.00025409972,0.954983,0.00021678038,0.00025787437,0.00018286589,0.000751246,0.0016588301,0.0046389094,0.00014153494,0.0007204159,0.035838727],"study_design_scores_gemma":[0.00000896664,0.00015509441,0.99538666,0.000020293888,0.000041663276,0.0000934537,0.00032612545,0.0021771232,0.00096704566,0.00022254739,0.0005861321,0.000014934962],"about_ca_topic_score_codex":0.0031299475,"about_ca_topic_score_gemma":0.005562305,"teacher_disagreement_score":0.0046264958,"about_ca_system_score_codex":0.00061823067,"about_ca_system_score_gemma":0.00054739293,"threshold_uncertainty_score":0.024467528},"labels":[],"label_agreement":null},{"id":"W2394841101","doi":"10.1109/saner.2016.56","title":"Defect Prediction: Accomplishments and Future Challenges","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":124,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software quality assurance; Software; Software quality; Software development; Software engineering; Quality (philosophy); Software metric; Data science; Software quality analyst; Key (lock); Field (mathematics); Prioritization; Risk analysis (engineering); Management science; Engineering; Computer security","score_opus":0.02160372313132453,"score_gpt":0.24602186263941647,"score_spread":0.22441813950809195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2394841101","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027908083,0.41626215,0.13193212,0.39863822,0.0061115357,0.00018548519,0.00082654133,0.001473464,0.016662449],"genre_scores_gemma":[0.28438574,0.42764542,0.23567238,0.022242863,0.014938404,0.00040761143,0.002365823,0.0005040061,0.011837714],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9873303,0.0056235134,0.0009244073,0.0021159821,0.0034869271,0.0005188745],"domain_scores_gemma":[0.8874106,0.0745417,0.0028931336,0.0060411184,0.024953969,0.0041593914],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.037134286,0.0021175495,0.0017538171,0.003534759,0.0018778669,0.008001696,0.0051415255,0.007643826,0.004487215],"category_scores_gemma":[0.06547109,0.0006564698,0.0013428247,0.0031569428,0.0047642915,0.021893293,0.0038857665,0.008090037,0.0028654907],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002124463,0.00042334004,0.010705409,0.0019366539,0.000141349,0.00022328459,0.0012069662,0.01356396,0.00095341855,0.067144774,0.058655985,0.8448324],"study_design_scores_gemma":[0.00008408532,0.0010186712,0.012788536,0.006695502,0.00023822254,0.0014532673,0.013715534,0.11475987,0.0034597912,0.42058074,0.4247062,0.00049970794],"about_ca_topic_score_codex":0.0062936526,"about_ca_topic_score_gemma":0.003969929,"teacher_disagreement_score":0.037134286,"about_ca_system_score_codex":0.0025760662,"about_ca_system_score_gemma":0.0038872617,"threshold_uncertainty_score":0.19638723},"labels":[],"label_agreement":null},{"id":"W2394850408","doi":"","title":"Taming a Tiger: software engineering in the era of big data & continuous development","year":2015,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Toronto Metropolitan University; Western University","funders":"","keywords":"Big data; Computer science; Software development; Software; Data science; Software engineering; Social software engineering; Quality (philosophy); Software construction; Data mining; Programming language","score_opus":0.04466692417880035,"score_gpt":0.25063165506620444,"score_spread":0.20596473088740408,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2394850408","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016183678,0.05994695,0.51222056,0.36440492,0.01692731,0.00014071315,0.00012098806,0.0027217772,0.027333137],"genre_scores_gemma":[0.24004872,0.06968285,0.5466656,0.07776654,0.009446165,0.00055379164,0.00064800767,0.0038551078,0.051333178],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98804075,0.0062042284,0.00048809298,0.00070534303,0.0039260965,0.00063555746],"domain_scores_gemma":[0.97310466,0.016958982,0.00077502156,0.002875049,0.0036274525,0.0026588575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025341261,0.00094876,0.0005245061,0.0021132969,0.0037258465,0.010449262,0.0019580547,0.003231374,0.004871202],"category_scores_gemma":[0.03273705,0.000560711,0.0008560536,0.0023997233,0.009406491,0.023284843,0.0077048102,0.00881415,0.00207995],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008939803,0.00012303091,0.0020653403,0.00068414706,0.00008229232,0.0006551009,0.011521501,0.004509037,0.0036591862,0.30092746,0.25760534,0.4180782],"study_design_scores_gemma":[0.000018104549,0.000107185995,0.0005486748,0.00085023785,0.000019816905,0.00036260174,0.006825422,0.004966483,0.0014479883,0.41267708,0.5720841,0.00009227065],"about_ca_topic_score_codex":0.0018616413,"about_ca_topic_score_gemma":0.0027148435,"teacher_disagreement_score":0.025341261,"about_ca_system_score_codex":0.0020819174,"about_ca_system_score_gemma":0.00573135,"threshold_uncertainty_score":0.13401896},"labels":[],"label_agreement":null},{"id":"W2395207450","doi":"","title":"Towards convenient management of software clone codes in practice: an integrated approach","year":2015,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Software maintenance; Software development; Cloning (programming); Software engineering; Computer science; Software; Software system; Software evolution; Software construction; Programming language; Biology","score_opus":0.026225326781531867,"score_gpt":0.27561971197360585,"score_spread":0.249394385192074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395207450","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02420411,0.00034036627,0.96681243,0.0008582101,0.00001791627,0.00052732445,0.000040942716,0.0044775284,0.002721292],"genre_scores_gemma":[0.08501444,0.00030708598,0.91160166,0.00012223894,0.000024234985,0.00032997903,0.00016209448,0.0004169391,0.0020211912],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97990143,0.00663195,0.0020146584,0.003183292,0.007631941,0.0006367804],"domain_scores_gemma":[0.94162464,0.017964777,0.008086705,0.01969548,0.010650646,0.0019776863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016422685,0.0010530966,0.001289327,0.005700229,0.0013896815,0.0073458645,0.0049161008,0.0024050944,0.0021252676],"category_scores_gemma":[0.05828917,0.0014911189,0.00096260133,0.0033952955,0.002816972,0.014166351,0.008638831,0.0025623022,0.0012101855],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028732343,0.00094306446,0.016802462,0.00080674276,0.000105119245,0.0007225012,0.018461416,0.0045433887,0.039789345,0.042301368,0.002741824,0.8724955],"study_design_scores_gemma":[0.00062821154,0.004283896,0.049853496,0.0033580267,0.0008146697,0.010609081,0.026845222,0.30922118,0.18179181,0.17076294,0.24100801,0.00082336296],"about_ca_topic_score_codex":0.0012496486,"about_ca_topic_score_gemma":0.0015838123,"teacher_disagreement_score":0.016422685,"about_ca_system_score_codex":0.0014598581,"about_ca_system_score_gemma":0.0043064156,"threshold_uncertainty_score":0.08685249},"labels":[],"label_agreement":null},{"id":"W2395258251","doi":"10.1145/2896982.2896983","title":"Examining the co-evolution relationship between simulink models and their test cases","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; General Motors of Canada","keywords":"MATLAB; Computer science; Test (biology); Relation (database); Value (mathematics); Empirical research; Work (physics); Machine learning; Data mining; Mathematics; Statistics; Engineering; Programming language; Ecology","score_opus":0.14142809561221106,"score_gpt":0.30497199885356197,"score_spread":0.1635439032413509,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395258251","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97715145,0.00015345532,0.01966645,0.00019378711,0.0000061291807,0.000085170956,0.0000648279,0.00011833153,0.0025603513],"genre_scores_gemma":[0.9834494,0.0000641699,0.015783425,0.000015514093,0.0000037492507,0.0000405614,0.00019518832,0.00003497862,0.0004130478],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9838409,0.0075546065,0.0008731433,0.0018238891,0.0051043807,0.00080301607],"domain_scores_gemma":[0.76743317,0.18094657,0.02274137,0.012485539,0.0147009995,0.0016923272],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0103604,0.00045674824,0.0003248738,0.0041515785,0.00069937715,0.0020618977,0.0012185182,0.00078406464,0.0010149054],"category_scores_gemma":[0.11430951,0.0005473201,0.00048294984,0.0028892376,0.0011236843,0.0022597257,0.0019974504,0.0012785206,0.00018901215],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029352683,0.0006481226,0.8127699,0.00020629005,0.0003033285,0.0019338123,0.008649814,0.053543687,0.012706403,0.0065796566,0.0003904531,0.10197496],"study_design_scores_gemma":[0.00006436578,0.0013663275,0.4739658,0.0001727643,0.00032999302,0.003619795,0.009087329,0.4658945,0.032764953,0.0053753546,0.0072245854,0.00013427882],"about_ca_topic_score_codex":0.0055424115,"about_ca_topic_score_gemma":0.0070559005,"teacher_disagreement_score":0.0103604,"about_ca_system_score_codex":0.0018431295,"about_ca_system_score_gemma":0.0012431454,"threshold_uncertainty_score":0.05479169},"labels":[],"label_agreement":null},{"id":"W2395314218","doi":"10.1145/2896995.2896999","title":"Measuring the principal of defect debt","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Technical debt; Software bug; Computer science; Principal (computer security); Software regression; Software; Schedule; Debt; Predictive power; Predictive modelling; Principal component analysis; Term (time); Econometrics; Software quality; Software development; Artificial intelligence; Business; Machine learning; Finance; Computer security; Economics","score_opus":0.03512044174850015,"score_gpt":0.25376626588298623,"score_spread":0.21864582413448608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395314218","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9488068,0.00041947316,0.04511864,0.00020108336,0.000026515869,0.000120480465,0.0018594732,0.00046281886,0.0029846902],"genre_scores_gemma":[0.98619556,0.00020458172,0.011125505,0.000014884751,0.000014551502,0.00004960728,0.0014025273,0.000026146632,0.0009666306],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9980057,0.0003323368,0.00029576357,0.0005741351,0.00065204117,0.00014004698],"domain_scores_gemma":[0.966605,0.01209104,0.012194456,0.002253002,0.005969468,0.0008869811],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028921773,0.0007701643,0.0004935323,0.0037612945,0.00029140868,0.0009563262,0.0005806222,0.0006354089,0.0010084906],"category_scores_gemma":[0.028157631,0.00040979573,0.0005025592,0.0029562495,0.00044826584,0.0015837558,0.00062893145,0.000873069,0.00046853258],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014649698,0.0002364755,0.85082346,0.0001704467,0.000113708404,0.00014305754,0.00034564687,0.042451907,0.003477013,0.0011241809,0.0014263353,0.09954124],"study_design_scores_gemma":[0.000012370365,0.00033101137,0.7005584,0.000063238156,0.000055437118,0.00037635668,0.00024352875,0.29050758,0.004029607,0.001856153,0.0019079451,0.000058377187],"about_ca_topic_score_codex":0.006727047,"about_ca_topic_score_gemma":0.007778056,"teacher_disagreement_score":0.006727047,"about_ca_system_score_codex":0.00092721794,"about_ca_system_score_gemma":0.00073050824,"threshold_uncertainty_score":0.0152955055},"labels":[],"label_agreement":null},{"id":"W2395760792","doi":"10.1145/2876441","title":"Understanding JavaScript Event-Based Interactions with Clematis","year":2016,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Intel Corporation","keywords":"Computer science; JavaScript; Program comprehension; Event (particle physics); Software engineering; Web application; Visualization; Asynchronous communication; Programming language; Human–computer interaction; Software; Software system; Artificial intelligence; World Wide Web","score_opus":0.1857356358688734,"score_gpt":0.3321283856873676,"score_spread":0.14639274981849418,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395760792","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4847833,0.00027082395,0.48570284,0.00049012434,0.000047230387,0.0002491143,0.0011714393,0.017321322,0.009963888],"genre_scores_gemma":[0.8314207,0.00015513942,0.16276968,0.00013630114,0.000026496287,0.00010627987,0.0015636642,0.00095897686,0.002862778],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99940825,0.00014250558,0.00003309064,0.00017619558,0.00019129716,0.000048590297],"domain_scores_gemma":[0.99667346,0.0021796736,0.00043018095,0.0003034464,0.00031812012,0.00009506097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00051570765,0.0006735823,0.0003200843,0.0009090802,0.0004241507,0.0017005209,0.0007442747,0.000936832,0.0022249678],"category_scores_gemma":[0.004944259,0.00035416498,0.0004208995,0.0003881984,0.00051728555,0.0018217535,0.00069891947,0.0009973927,0.00062573183],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016735003,0.0012107008,0.07766018,0.0012643597,0.00017008538,0.004091684,0.011615396,0.13093407,0.33316782,0.020269929,0.011042559,0.4068998],"study_design_scores_gemma":[0.000038760438,0.00019805081,0.04795338,0.00007758341,0.00004641796,0.00067126297,0.0005924123,0.87245613,0.055986352,0.008432076,0.013479848,0.000067835834],"about_ca_topic_score_codex":0.005272829,"about_ca_topic_score_gemma":0.006575718,"teacher_disagreement_score":0.005272829,"about_ca_system_score_codex":0.00063431315,"about_ca_system_score_gemma":0.0005535263,"threshold_uncertainty_score":0.010484278},"labels":[],"label_agreement":null},{"id":"W2395791174","doi":"10.1145/2889160.2889244","title":"CoRReCT","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":87,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science","score_opus":0.02493018398683174,"score_gpt":0.28335076242512197,"score_spread":0.25842057843829025,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395791174","genre_codex":"other","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011286918,0.00420291,0.17481482,0.020934997,0.046420146,0.0020728433,0.028801307,0.06982886,0.6416372],"genre_scores_gemma":[0.06642838,0.0033340077,0.10398244,0.0055096345,0.0057446132,0.00082421495,0.03392832,0.022268029,0.75798035],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99212575,0.0010809938,0.0006138198,0.0013302742,0.0043273037,0.0005217846],"domain_scores_gemma":[0.96015716,0.003985574,0.0015687532,0.00975809,0.023268342,0.0012619551],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043071127,0.0017114981,0.0009907958,0.005241519,0.002960082,0.00569904,0.0024430451,0.002570703,0.34139058],"category_scores_gemma":[0.038711313,0.000649902,0.0012505131,0.0030788353,0.001065566,0.0051837633,0.0042259353,0.002227408,0.31761146],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010541721,0.000038600392,0.0018073783,0.00056605326,0.000035183748,0.00036023112,0.00033669834,0.00050131034,0.0025977483,0.010759054,0.69589144,0.28700095],"study_design_scores_gemma":[0.000019899657,0.000019273453,0.0009530422,0.00018073294,0.000018435936,0.0003528902,0.00017627944,0.000876171,0.0022554547,0.004406448,0.99070907,0.000032254375],"about_ca_topic_score_codex":0.0046774605,"about_ca_topic_score_gemma":0.0064325384,"teacher_disagreement_score":0.34139058,"about_ca_system_score_codex":0.0015375583,"about_ca_system_score_gemma":0.005086319,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2396100201","doi":"10.1109/saner.2016.113","title":"The Impact of Human Discussions on Just-in-Time Quality Assurance: An Empirical Study on OpenStack and Eclipse","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Eclipse; Source lines of code; Variety (cybernetics); Process (computing); Quality (philosophy); Recall; Relation (database); Logistic regression; Code review; Feeling; Empirical research; Data science; Machine learning; Data mining; Artificial intelligence; Software quality; Psychology; Cognitive psychology; Programming language; Statistics; Software; Software development; Social psychology","score_opus":0.08834997096515398,"score_gpt":0.45344485873854345,"score_spread":0.3650948877733895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396100201","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99920386,0.00008003996,0.00022381815,0.0000677325,0.0000027539868,0.000025234953,0.00006434545,0.000008882809,0.00032339853],"genre_scores_gemma":[0.99897754,0.000052978434,0.00033233542,0.000039125964,0.000010329638,0.00004559928,0.00015936368,0.000014840247,0.00036781194],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9830329,0.010795,0.0013514464,0.0015251514,0.0025984335,0.0006971733],"domain_scores_gemma":[0.38744292,0.51557535,0.06637719,0.009343638,0.015208111,0.006052763],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025525011,0.00041725455,0.00045188985,0.00212042,0.0010086708,0.0020743792,0.0011205064,0.0011116136,0.0021802855],"category_scores_gemma":[0.18230677,0.00039195508,0.00059670064,0.0016346594,0.0015898924,0.0033387505,0.0013583059,0.0018230745,0.00063737226],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010506212,0.0019576696,0.9460491,0.00028618926,0.00015384164,0.00033339785,0.024842048,0.0009523778,0.0007804777,0.00024235925,0.0006697072,0.022682322],"study_design_scores_gemma":[0.000055117813,0.0015826725,0.97766113,0.00010144332,0.00007836481,0.0002629945,0.012148984,0.0054546166,0.00074084336,0.0002898568,0.0015630213,0.00006093345],"about_ca_topic_score_codex":0.0052284426,"about_ca_topic_score_gemma":0.0051849713,"teacher_disagreement_score":0.025525011,"about_ca_system_score_codex":0.0015671103,"about_ca_system_score_gemma":0.0011030188,"threshold_uncertainty_score":0.13499081},"labels":[],"label_agreement":null},{"id":"W2396152402","doi":"","title":"Temporal Software Change Prediction Using Neural Networks.","year":2007,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Artificial neural network; Machine learning; Artificial intelligence; Data mining; Data science; Programming language","score_opus":0.025687410420802685,"score_gpt":0.2537456631959293,"score_spread":0.2280582527751266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396152402","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47145706,0.0073942603,0.49470788,0.0024365191,0.0006883606,0.00026873284,0.004887452,0.006252977,0.0119068],"genre_scores_gemma":[0.93610406,0.0007608638,0.056934033,0.0001509458,0.00013334063,0.00011306665,0.002879008,0.00008440798,0.00284028],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994517,0.000117508825,0.000039654635,0.00018879508,0.00013428277,0.000068027315],"domain_scores_gemma":[0.99670166,0.0020086004,0.00045794624,0.00022337011,0.00050746644,0.00010094451],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011499642,0.0006488263,0.00050517695,0.0020980102,0.00031762625,0.0007257492,0.00095747184,0.0009358306,0.0017350898],"category_scores_gemma":[0.008192508,0.00031283108,0.0005276285,0.0017973584,0.00028158512,0.0014740889,0.00050055556,0.0011782668,0.00057028554],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008001705,0.00046415013,0.056311753,0.00027831364,0.00040775354,0.00030035683,0.000112042544,0.4490994,0.0030087503,0.0032243673,0.011848062,0.47414488],"study_design_scores_gemma":[0.00000763303,0.000017444921,0.0027825325,0.00000983334,0.000020108493,0.000018181763,0.000010509429,0.9935597,0.0006402126,0.0024696346,0.00045951255,0.0000046318128],"about_ca_topic_score_codex":0.01711713,"about_ca_topic_score_gemma":0.027621511,"teacher_disagreement_score":0.01711713,"about_ca_system_score_codex":0.0009906711,"about_ca_system_score_gemma":0.00046614656,"threshold_uncertainty_score":0.034035027},"labels":[],"label_agreement":null},{"id":"W2396272922","doi":"","title":"Pilot study of collective decision-making in the code review process","year":2015,"lang":"en","type":"article","venue":"Institutional Repositories DataBase (IRDB)","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Voting; Computer science; Simple (philosophy); Majority rule; Code review; Process (computing); Code (set theory); Core (optical fiber); Point (geometry); Computer security; Software; Politics; Artificial intelligence; Political science; Mathematics; Law; Static program analysis; Programming language; Software development; Set (abstract data type)","score_opus":0.06962856520644845,"score_gpt":0.3610349439929956,"score_spread":0.2914063787865472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396272922","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99349916,0.000015238734,0.0016609773,0.00020134085,0.000033170854,0.0015708213,0.000051219722,0.000053821797,0.0029143207],"genre_scores_gemma":[0.990297,0.00002336355,0.006111358,0.00013090954,0.000024733843,0.0019860484,0.000085168715,0.000018701896,0.0013225501],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9799917,0.015144859,0.00081214134,0.0011867101,0.0015869474,0.0012776654],"domain_scores_gemma":[0.6458019,0.29652166,0.010650036,0.014092979,0.01727678,0.015656546],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03378636,0.00057867356,0.0007743067,0.0013229633,0.004868473,0.002977196,0.002016406,0.0018852046,0.008365757],"category_scores_gemma":[0.12972948,0.00082771824,0.00060129096,0.0012012979,0.0020838461,0.002814993,0.0025881012,0.0027780528,0.0013733879],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.028873416,0.25340083,0.1476919,0.0021711902,0.0004485363,0.0035143138,0.29423282,0.010579962,0.018284231,0.008518457,0.008753769,0.22353065],"study_design_scores_gemma":[0.020406809,0.22918826,0.3484606,0.0006961741,0.0007787034,0.00077156007,0.27334353,0.06673373,0.017762685,0.011931795,0.029177671,0.00074849406],"about_ca_topic_score_codex":0.006647939,"about_ca_topic_score_gemma":0.010010607,"teacher_disagreement_score":0.96621364,"about_ca_system_score_codex":0.003026046,"about_ca_system_score_gemma":0.0089063775,"threshold_uncertainty_score":0.1786815},"labels":[],"label_agreement":null},{"id":"W2396354602","doi":"","title":"How should we read and analyze bug reports: an interactive visualization using extractive summaries and topic evolution","year":2015,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Visualization; Computer science; Software bug; Task (project management); Program comprehension; Comprehension; Software; Software visualization; Data visualization; Creative visualization; Software engineering; Data science; World Wide Web; Human–computer interaction; Software development; Data mining; Software system; Engineering; Systems engineering; Programming language; Component-based software engineering","score_opus":0.04506052755335299,"score_gpt":0.30098692023822443,"score_spread":0.25592639268487144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396354602","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25653172,0.0037366292,0.681477,0.004751051,0.0005142755,0.0007487772,0.0034388953,0.04236083,0.0064408095],"genre_scores_gemma":[0.40503672,0.0020560138,0.5858849,0.00030219916,0.00029843385,0.00039660002,0.0018884164,0.001619995,0.002516622],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989573,0.00047543363,0.000103738086,0.00019581716,0.00021096595,0.000056662353],"domain_scores_gemma":[0.9903106,0.005854942,0.001232228,0.0009520721,0.0012766927,0.0003735039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002833641,0.001781166,0.000714011,0.0037505697,0.0005530717,0.0033177463,0.00084391993,0.001104029,0.0031611954],"category_scores_gemma":[0.016936518,0.0005207555,0.00072715007,0.001771274,0.00036134105,0.004664562,0.0013656913,0.0009570376,0.0009167656],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016437687,0.00040065887,0.018555526,0.003384605,0.00033084588,0.0013890001,0.034476403,0.0059812632,0.10444156,0.0039002271,0.036185946,0.7893101],"study_design_scores_gemma":[0.0011396025,0.0038036236,0.12261071,0.0032215165,0.0017589697,0.0066414657,0.029888287,0.3211592,0.14625315,0.032903157,0.32926446,0.0013558519],"about_ca_topic_score_codex":0.0010993812,"about_ca_topic_score_gemma":0.0013409212,"teacher_disagreement_score":0.0037505697,"about_ca_system_score_codex":0.00027431786,"about_ca_system_score_gemma":0.00042768545,"threshold_uncertainty_score":0.014985859},"labels":[],"label_agreement":null},{"id":"W2396547321","doi":"","title":"An empirical study on change recommendation","year":2015,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Programmer; Computer science; Reuse; Ranking (information retrieval); Focus (optics); Repetition (rhetorical device); Sensitivity (control systems); Empirical research; Recommender system; Software engineering; Information retrieval; Programming language; Statistics","score_opus":0.0910226896818921,"score_gpt":0.34032712228105266,"score_spread":0.24930443259916057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396547321","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98815155,0.00084418635,0.0040410017,0.00088278233,0.000027528196,0.00018933526,0.0013759946,0.00010233918,0.004385213],"genre_scores_gemma":[0.99104184,0.00037784304,0.0053266804,0.000245854,0.00003533061,0.0001308407,0.0017974916,0.00004054772,0.0010036216],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97030365,0.016720345,0.003156637,0.002835584,0.006176619,0.00080715626],"domain_scores_gemma":[0.33859798,0.58231765,0.034328684,0.020616647,0.021723034,0.0024159905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025796624,0.00039149524,0.0005434637,0.0040516434,0.0012526434,0.0025811263,0.0017613481,0.0015568228,0.0047463058],"category_scores_gemma":[0.2950792,0.00049288775,0.00048510058,0.0071860105,0.0012719345,0.005645221,0.0009165621,0.0024166508,0.0012399046],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005955784,0.0017411736,0.91653717,0.0007772442,0.00028871887,0.00042718314,0.0045065586,0.002267909,0.0007645719,0.001535984,0.0042930543,0.066264935],"study_design_scores_gemma":[0.00023231359,0.002252961,0.8689242,0.00056043535,0.00034363172,0.002883214,0.016145645,0.075874515,0.0036065513,0.0027704297,0.026244378,0.00016172892],"about_ca_topic_score_codex":0.007007511,"about_ca_topic_score_gemma":0.006684636,"teacher_disagreement_score":0.025796624,"about_ca_system_score_codex":0.0011914854,"about_ca_system_score_gemma":0.0013121889,"threshold_uncertainty_score":0.13642722},"labels":[],"label_agreement":null},{"id":"W2396714452","doi":"10.1142/s0129054105003327","title":"REPRESENTATION OF SEMIAUTOMATA BY CANONICAL WORDS AND EQUIVALENCES","year":2005,"lang":"en","type":"article","venue":"International Journal of Foundations of Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; University of Waterloo","funders":"","keywords":"Prefix; Rewriting; Mathematics; Canonical form; Equivalence (formal languages); State (computer science); Word (group theory); Set (abstract data type); Discrete mathematics; Computer science; Algorithm; Pure mathematics; Programming language","score_opus":0.022387374634667168,"score_gpt":0.35223913773278087,"score_spread":0.3298517630981137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396714452","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047054276,0.000277846,0.94115967,0.0001571626,0.00007209351,0.000081075785,0.00014592193,0.00072441803,0.010327501],"genre_scores_gemma":[0.55626225,0.00041731182,0.4326839,0.00013692878,0.00015874568,0.00033225576,0.0004792797,0.0004624122,0.009066979],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9975107,0.00064337725,0.00027631863,0.0006611303,0.0006153891,0.00029312872],"domain_scores_gemma":[0.99705076,0.0012214466,0.00038525657,0.0006670252,0.00051830715,0.000157198],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001639126,0.0005346039,0.00052806776,0.0015412185,0.0009334334,0.002750405,0.0013425268,0.00093633943,0.0049450826],"category_scores_gemma":[0.0052021747,0.0005261362,0.0011790555,0.0011384045,0.003046905,0.0054307794,0.002297304,0.0016538122,0.0010957463],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00001766839,0.000026054777,0.00018460852,0.000038886497,0.0000068172303,0.0001248608,0.00043864534,0.0050159814,0.0024654137,0.97575676,0.00026853263,0.015655879],"study_design_scores_gemma":[0.000015008623,0.000042158124,0.00008885521,0.000026807738,0.00001523698,0.00012932031,0.00009902197,0.035428345,0.0044476315,0.9501135,0.0095744375,0.000019639707],"about_ca_topic_score_codex":0.0008307145,"about_ca_topic_score_gemma":0.00084023626,"teacher_disagreement_score":0.0049450826,"about_ca_system_score_codex":0.0009713055,"about_ca_system_score_gemma":0.0008548753,"threshold_uncertainty_score":0.016542912},"labels":[],"label_agreement":null},{"id":"W2397274409","doi":"10.1002/smr.1791","title":"A Simple, Efficient, Context‐sensitive Approach for Code Completion","year":2016,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Code (set theory); Context (archaeology); Field (mathematics); Security token; Source code; Artificial intelligence; Set (abstract data type); Machine learning; Natural language processing; Programming language; Computer security","score_opus":0.022473119251122036,"score_gpt":0.28199678673051376,"score_spread":0.2595236674793917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397274409","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025557088,0.0003281367,0.949944,0.00019910927,0.000092134076,0.00040977562,0.0002902653,0.021793906,0.0013856622],"genre_scores_gemma":[0.17879687,0.00015121402,0.8146704,0.0001660248,0.000059458627,0.0002655562,0.0009969224,0.0014428798,0.0034506924],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9933687,0.0012135186,0.00040544028,0.0013910963,0.0033580363,0.0002632733],"domain_scores_gemma":[0.98228097,0.004576369,0.0019350963,0.005719821,0.0047636544,0.00072405604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034184516,0.0016333697,0.0013039259,0.0047174892,0.0011334216,0.0015756404,0.0032080517,0.0012952134,0.0041235103],"category_scores_gemma":[0.02610914,0.0010239576,0.0015850208,0.0025077667,0.00137461,0.003132686,0.0035247544,0.0029812155,0.0027790163],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048604602,0.00050009286,0.007720383,0.00045859165,0.00011012572,0.00032401906,0.0011025209,0.042902738,0.036562838,0.012429175,0.013194737,0.8842087],"study_design_scores_gemma":[0.00007898806,0.0003443394,0.0022584193,0.000067505265,0.00007963886,0.00042856715,0.00027398413,0.92028517,0.041695066,0.012112653,0.022246817,0.00012891478],"about_ca_topic_score_codex":0.008022575,"about_ca_topic_score_gemma":0.008690031,"teacher_disagreement_score":0.008022575,"about_ca_system_score_codex":0.001141713,"about_ca_system_score_gemma":0.005294193,"threshold_uncertainty_score":0.018078685},"labels":[],"label_agreement":null},{"id":"W2397351127","doi":"10.1016/b978-0-12-800161-5.00004-9","title":"Multiobjective Optimization for Software Refactoring and Evolution","year":2014,"lang":"en","type":"book-chapter","venue":"Advances in computers","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Code refactoring; Computer science; Process (computing); Software maintenance; Identification (biology); Software; Software engineering; Automation; Code (set theory); Adaptation (eye); Programming language; Software system; Engineering","score_opus":0.012424789861069601,"score_gpt":0.26298334516133953,"score_spread":0.25055855530026994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397351127","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024266778,0.018802794,0.9295088,0.0005342692,0.00052075885,0.00003967484,0.000100838624,0.0003546836,0.04771143],"genre_scores_gemma":[0.07886751,0.024782194,0.7711165,0.00045736716,0.0006561093,0.0004033527,0.0004918621,0.0008636403,0.122361474],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99965084,0.000082168655,0.000015441567,0.00004466018,0.00018976259,0.000017152412],"domain_scores_gemma":[0.999835,0.00008960196,0.000014308517,0.000020981195,0.000033346343,0.0000066934817],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00044024453,0.0016650043,0.00095568364,0.0007248234,0.00030747123,0.00090152276,0.0012348786,0.0010709914,0.008020093],"category_scores_gemma":[0.00074416865,0.0004610008,0.0010280546,0.0014840856,0.00049240596,0.0008403107,0.00096201675,0.0018635083,0.0025456364],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002290444,0.000072820636,0.00011308882,0.000622999,0.00009530165,0.00006299505,0.000069833375,0.33333087,0.005515542,0.09790383,0.022426365,0.53976345],"study_design_scores_gemma":[0.000018289418,0.00008021676,0.00040348235,0.00041090316,0.00005001521,0.00016317406,0.00003374652,0.7119424,0.004202469,0.15526786,0.12738119,0.000046213823],"about_ca_topic_score_codex":0.0010205957,"about_ca_topic_score_gemma":0.001614415,"teacher_disagreement_score":0.008020093,"about_ca_system_score_codex":0.0008230162,"about_ca_system_score_gemma":0.00042309734,"threshold_uncertainty_score":0.026829898},"labels":[],"label_agreement":null},{"id":"W2397486511","doi":"10.7287/peerj.preprints.1260","title":"Comments on \"Researcher bias: The use of machine learning in software defect prediction\"","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; McGill University","funders":"","keywords":"Metric (unit); Association (psychology); Construct (python library); Computer science; Reuse; Group (periodic table); Machine learning; Predictive modelling; Artificial intelligence; Software; Selection (genetic algorithm); Data mining; Data science; Econometrics; Psychology; Mathematics; Engineering; Operations management","score_opus":0.10402786334163419,"score_gpt":0.2983163283872558,"score_spread":0.1942884650456216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397486511","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005882191,0.00080969924,0.00066891767,0.97953963,0.017433418,0.000020460448,0.00015232817,0.00009167969,0.00069566513],"genre_scores_gemma":[0.0058806525,0.0009514168,0.0009192139,0.9703944,0.019579913,0.00007790981,0.00006698527,0.00010557605,0.0020239938],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9582209,0.017250098,0.0045869467,0.004388088,0.013955272,0.0015987214],"domain_scores_gemma":[0.67526144,0.22653627,0.014850997,0.008160314,0.068440005,0.0067510614],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.048404817,0.002000958,0.0014360285,0.0024092612,0.0059588565,0.006100253,0.0071121296,0.033039514,0.0062715933],"category_scores_gemma":[0.24541762,0.0012621022,0.00218852,0.0032127746,0.007923996,0.007621877,0.0044431835,0.03678506,0.005506236],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044150696,0.00001871588,0.00089339283,0.00011830312,0.00002224398,0.00027631264,0.0008473462,0.0001597917,0.00018090135,0.0019012607,0.99064165,0.004895996],"study_design_scores_gemma":[0.00010678522,0.00014192323,0.003949596,0.001459583,0.00009390068,0.0011240627,0.0056331595,0.001512394,0.001565065,0.008674436,0.9754169,0.00032222152],"about_ca_topic_score_codex":0.017158234,"about_ca_topic_score_gemma":0.014183477,"teacher_disagreement_score":0.9515952,"about_ca_system_score_codex":0.0055414443,"about_ca_system_score_gemma":0.008826547,"threshold_uncertainty_score":0.25599217},"labels":[],"label_agreement":null},{"id":"W2398184019","doi":"","title":"Understanding Expert Perception in Software Estimation Effort: a Cognitive Approach Using Software Chunks.","year":2014,"lang":"en","type":"article","venue":"Cognitive Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Analogy; Computer science; Intuition; Categorization; Software development; Software; Perception; Artificial intelligence; Cognition; Estimation; Machine learning; Cognitive science; Psychology; Engineering; Systems engineering; Programming language","score_opus":0.10281975489215141,"score_gpt":0.328239075028085,"score_spread":0.22541932013593363,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398184019","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9065109,0.0005422594,0.07762723,0.0010584759,0.000019180154,0.0000704394,0.000059092275,0.000059212834,0.01405322],"genre_scores_gemma":[0.9926086,0.00008804831,0.006861076,0.000056915207,0.0000052187224,0.000017555903,0.000022160384,0.000006865176,0.00033357835],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99856025,0.00053259026,0.00005983763,0.00029413195,0.00044072865,0.00011244063],"domain_scores_gemma":[0.9815011,0.01276081,0.0022125456,0.000965151,0.001864023,0.0006964269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031166049,0.00035701675,0.0002352042,0.002145257,0.00065228384,0.0030477261,0.0006992129,0.0010126855,0.002133981],"category_scores_gemma":[0.030546669,0.0004349505,0.00043556225,0.0007376819,0.0020711618,0.005425991,0.0019792118,0.0009490082,0.00016595285],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070121896,0.000710481,0.28029513,0.00079135306,0.00025944645,0.0009649654,0.22252722,0.013937993,0.03609906,0.07766665,0.0039237333,0.36212274],"study_design_scores_gemma":[0.00009118935,0.00056740624,0.49170274,0.00048543198,0.00021985742,0.0011481547,0.11996626,0.19248104,0.008159931,0.17287388,0.011955407,0.0003486805],"about_ca_topic_score_codex":0.015071524,"about_ca_topic_score_gemma":0.014346542,"teacher_disagreement_score":0.015071524,"about_ca_system_score_codex":0.0018057877,"about_ca_system_score_gemma":0.0009177877,"threshold_uncertainty_score":0.029967546},"labels":[],"label_agreement":null},{"id":"W2399079435","doi":"10.13140/rg.2.1.3606.9201","title":"Recommending Relevant Sections from a Webpage about Programming Errors and Exceptions","year":2015,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Web page; Context (archaeology); World Wide Web; Information retrieval; Software; Static web page; The Internet; Page view; Precision and recall; Web development; Programming language","score_opus":0.10154128332210781,"score_gpt":0.21505407761076376,"score_spread":0.11351279428865595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399079435","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5288411,0.0075134076,0.39251933,0.001378052,0.00062588206,0.00159016,0.007884989,0.041202374,0.018444797],"genre_scores_gemma":[0.42364043,0.0026771466,0.54690546,0.00028529187,0.000335458,0.00032283945,0.010647088,0.0008714783,0.014314805],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99961346,0.00005884465,0.000038127404,0.00011372064,0.00014824544,0.000027638807],"domain_scores_gemma":[0.9973877,0.0010518649,0.0003045702,0.00024709932,0.00085101876,0.00015770152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00046715632,0.0011884262,0.0007524782,0.0051043252,0.00064589723,0.0010043418,0.00066620245,0.001039937,0.0026546877],"category_scores_gemma":[0.004648813,0.00045928752,0.00072998455,0.0021805915,0.00018741265,0.0012072392,0.00037296236,0.00068711344,0.0033772045],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005003837,0.0006884451,0.04199126,0.001235258,0.00017531675,0.00078571663,0.00069575594,0.0046815313,0.05106794,0.0008169655,0.03604043,0.861321],"study_design_scores_gemma":[0.0002235582,0.0020315463,0.18329568,0.0011682927,0.0016716557,0.0063082334,0.0031692453,0.46039885,0.1780218,0.009125649,0.1542144,0.0003711444],"about_ca_topic_score_codex":0.0048253485,"about_ca_topic_score_gemma":0.014481432,"teacher_disagreement_score":0.0051043252,"about_ca_system_score_codex":0.0002585033,"about_ca_system_score_gemma":0.0011315413,"threshold_uncertainty_score":0.00959456},"labels":[],"label_agreement":null},{"id":"W2400579873","doi":"10.1145/2889160.2889168","title":"JDeodorant","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Eclipse; Computer science; Java; Programming language; Software engineering; Plug-in; Code (set theory); Software maintenance; Software evolution; Software; Software system; Software construction","score_opus":0.016782784639996346,"score_gpt":0.2514172523715865,"score_spread":0.23463446773159016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2400579873","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007669811,0.002020831,0.21439907,0.000709823,0.00080738513,0.0003377869,0.014830421,0.70785296,0.051371988],"genre_scores_gemma":[0.12008916,0.0034905747,0.36641595,0.003106721,0.00041213338,0.0016413111,0.07638901,0.29030538,0.13814974],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984465,0.00013006706,0.00012150435,0.00037142783,0.000770839,0.0001596534],"domain_scores_gemma":[0.9973201,0.0010588049,0.00019725911,0.0007059325,0.0005314139,0.00018649065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018536495,0.0013257401,0.0010350671,0.0019345842,0.0006518821,0.002658349,0.0030637984,0.0017169318,0.034304023],"category_scores_gemma":[0.008442896,0.0015659558,0.0014366303,0.00083722477,0.0006976959,0.00446111,0.0036912358,0.0023950143,0.03037477],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011752292,0.00022684362,0.005120727,0.0022135004,0.0001750164,0.00092351565,0.00069788785,0.0015708378,0.0332943,0.01779339,0.6310399,0.30576882],"study_design_scores_gemma":[0.00027085596,0.000094781804,0.003106444,0.0002871735,0.00007000893,0.0015262546,0.00006140063,0.011626756,0.026984671,0.008162767,0.94765466,0.00015414198],"about_ca_topic_score_codex":0.0023551993,"about_ca_topic_score_gemma":0.003970908,"teacher_disagreement_score":0.034304023,"about_ca_system_score_codex":0.00062092836,"about_ca_system_score_gemma":0.0011054019,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2401028838","doi":"10.1145/2901739.2903493","title":"Judging a commit by its cover","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Commit; Computer science; Proxy (statistics); Code (set theory); Source code; Entropy (arrow of time); Programming language; Computer security; Database; Machine learning","score_opus":0.014958593381296096,"score_gpt":0.2481929661559607,"score_spread":0.2332343727746646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401028838","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94950277,0.000533958,0.036130093,0.00064389664,0.00022458816,0.00017337441,0.002815084,0.0015621819,0.008414125],"genre_scores_gemma":[0.9792572,0.00012114024,0.015549243,0.00008783362,0.00011982385,0.00008641806,0.0031552254,0.0002520757,0.0013710852],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99364156,0.0010619856,0.00075837283,0.0007355963,0.003430064,0.00037249364],"domain_scores_gemma":[0.9169647,0.045977015,0.010292063,0.0062838285,0.017718092,0.0027643587],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.005414511,0.0006104963,0.0005904495,0.005384752,0.0008105858,0.002050405,0.00059781375,0.0010122258,0.0023115824],"category_scores_gemma":[0.08777857,0.00031261554,0.00032979724,0.0027662579,0.00076659553,0.0027523234,0.0021867529,0.0010865751,0.0013077529],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013954954,0.00016597066,0.6725234,0.000710891,0.00023990158,0.00049416703,0.0044323383,0.008453177,0.026504774,0.0039993394,0.015859688,0.26522082],"study_design_scores_gemma":[0.00006231189,0.0008554997,0.7695021,0.0002607154,0.00017143117,0.0011563712,0.0046082116,0.16374283,0.023183128,0.010412722,0.025822256,0.00022246984],"about_ca_topic_score_codex":0.002240658,"about_ca_topic_score_gemma":0.004388489,"teacher_disagreement_score":0.9945855,"about_ca_system_score_codex":0.0005525137,"about_ca_system_score_gemma":0.00076670107,"threshold_uncertainty_score":0.028635025},"labels":[],"label_agreement":null},{"id":"W2401290433","doi":"10.1145/2901739.2901770","title":"Mining duplicate questions in stack overflow","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":123,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Reputation; Recall rate; Information retrieval; Precision and recall; Data mining; Recall; Stack (abstract data type); Data science; World Wide Web; Artificial intelligence; Programming language","score_opus":0.022784264838843913,"score_gpt":0.27972781170248784,"score_spread":0.2569435468636439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401290433","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.87691367,0.0031307149,0.10514316,0.00087197335,0.0002148339,0.00083922816,0.005495946,0.0038865162,0.0035039778],"genre_scores_gemma":[0.84806496,0.00080178346,0.13089195,0.00066057843,0.0002622495,0.00055894424,0.013705607,0.000400705,0.004653338],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9859735,0.003572185,0.0020750465,0.0024790214,0.004996622,0.0009037326],"domain_scores_gemma":[0.92868716,0.044416685,0.009548705,0.004336383,0.0115073575,0.0015037915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0085950615,0.0013637142,0.0014365851,0.011736788,0.0019351495,0.002399173,0.0021747581,0.003220636,0.0013328037],"category_scores_gemma":[0.061345275,0.00064291677,0.0010483869,0.004976668,0.0011540821,0.005091816,0.003421395,0.0012853104,0.00088160863],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014980558,0.0008691598,0.39285445,0.0037023246,0.00040179098,0.010164745,0.021722568,0.008825946,0.044542067,0.008522521,0.028508984,0.47838733],"study_design_scores_gemma":[0.0002069782,0.0013449147,0.31636912,0.0013692399,0.0011287207,0.024569064,0.023610651,0.33229238,0.14149843,0.03122321,0.12587501,0.00051232654],"about_ca_topic_score_codex":0.004195483,"about_ca_topic_score_gemma":0.0039410186,"teacher_disagreement_score":0.011736788,"about_ca_system_score_codex":0.001283797,"about_ca_system_score_gemma":0.0022740003,"threshold_uncertainty_score":0.045455575},"labels":[],"label_agreement":null},{"id":"W2401620832","doi":"10.1145/2901739.2903496","title":"The relationship between commit message detail and defect proneness in Java projects on GitHub","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Commit; Computer science; Explanatory power; Java; Code (set theory); Predictive power; Programming language; Database; Set (abstract data type)","score_opus":0.06137851503349717,"score_gpt":0.2916415823247079,"score_spread":0.2302630672912107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401620832","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99904674,0.000074643736,0.00042664498,0.000045700093,0.0000029545406,0.0000056164677,0.00018652066,0.00007917517,0.00013198235],"genre_scores_gemma":[0.99844044,0.000044661912,0.00044259147,0.000007781182,0.000007035507,0.0000061184337,0.0008068819,0.000029734085,0.00021471293],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99817467,0.00046011098,0.00022807643,0.00043336442,0.0005098855,0.00019388077],"domain_scores_gemma":[0.9014839,0.06656856,0.019288896,0.005404146,0.0051812953,0.002073096],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004238392,0.00057377684,0.0003618026,0.0032926428,0.0003154414,0.0012436446,0.0004926489,0.0007189147,0.0006402478],"category_scores_gemma":[0.046141896,0.00039661926,0.000458072,0.002212167,0.0006507982,0.0018560253,0.0010226873,0.0010842574,0.0003120328],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015642456,0.00012281514,0.9678284,0.000047988022,0.00010748269,0.00016396985,0.00054957945,0.015964948,0.0007043058,0.00013092517,0.00053339865,0.013689838],"study_design_scores_gemma":[0.000005774968,0.00018562711,0.9246562,0.000025546584,0.000047805064,0.00025120616,0.00030800622,0.07301036,0.00083757675,0.00032813664,0.00031975543,0.000023938175],"about_ca_topic_score_codex":0.0068839765,"about_ca_topic_score_gemma":0.009957419,"teacher_disagreement_score":0.99576163,"about_ca_system_score_codex":0.00061336433,"about_ca_system_score_gemma":0.0003818983,"threshold_uncertainty_score":0.022415042},"labels":[],"label_agreement":null},{"id":"W2401653830","doi":"10.1109/saner.2016.53","title":"An Empirical Study on Ranking Change Recommendations Retrieved Using Code Similarity","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Ranking (information retrieval); Computer science; Snippet; Rank (graph theory); Basis (linear algebra); Similarity (geometry); Code (set theory); Information retrieval; Change detection; Data mining; Artificial intelligence; Machine learning; Mathematics","score_opus":0.23158393243564904,"score_gpt":0.4271928547921154,"score_spread":0.19560892235646635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401653830","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9897968,0.0013407322,0.0053878683,0.00022858116,0.000030423504,0.00020425813,0.0011416471,0.0002934181,0.0015762523],"genre_scores_gemma":[0.9784303,0.0004229804,0.017698571,0.000060816965,0.00004715504,0.000104417486,0.0024460999,0.00005013796,0.0007395572],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98102796,0.009093732,0.002105425,0.0016509461,0.005605361,0.0005165077],"domain_scores_gemma":[0.7291464,0.22648638,0.013330281,0.011489212,0.018009312,0.0015383769],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013241313,0.00069520035,0.00093440997,0.007072527,0.00089715427,0.0019581618,0.0014757183,0.0012546112,0.0017216835],"category_scores_gemma":[0.13668172,0.00035681837,0.0005969424,0.007147285,0.000786128,0.0039064502,0.0006292076,0.0011320312,0.0007524151],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00264925,0.0032748585,0.6750385,0.0023282098,0.0009220602,0.0005192778,0.0020884948,0.010426358,0.010318228,0.0010147564,0.0069906395,0.2844294],"study_design_scores_gemma":[0.0005656577,0.007869405,0.6533661,0.00040504572,0.0008075502,0.0027857402,0.005163496,0.3009816,0.016069729,0.0013220434,0.0104274,0.00023626107],"about_ca_topic_score_codex":0.007750207,"about_ca_topic_score_gemma":0.013300269,"teacher_disagreement_score":0.013241313,"about_ca_system_score_codex":0.0008242339,"about_ca_system_score_gemma":0.0009801997,"threshold_uncertainty_score":0.07002753},"labels":[],"label_agreement":null},{"id":"W2402006257","doi":"10.1145/2901739.2903498","title":"The dispersion of build maintenance activity across maven lifecycle phases","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Application lifecycle management; Software maintenance; Computer science; Deliverable; Software engineering; Overhead (engineering); Software; Compiler; Bridge (graph theory); Systems engineering; Software development; Operating system; Engineering","score_opus":0.013747367605134076,"score_gpt":0.29105516227494527,"score_spread":0.2773077946698112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402006257","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97446537,0.0012668972,0.015404685,0.00028093637,0.000021048256,0.00008535103,0.0016269424,0.0015364995,0.0053123957],"genre_scores_gemma":[0.9768794,0.00074060116,0.013499576,0.00008976341,0.000024733457,0.00015639662,0.0057774284,0.0008664601,0.0019656823],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9913244,0.0019609726,0.0011626835,0.0014744672,0.0034703836,0.0006070675],"domain_scores_gemma":[0.9298475,0.032331306,0.014331293,0.01279635,0.009595203,0.0010983055],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062756212,0.00050241046,0.00052476086,0.012659769,0.0008195727,0.0028695236,0.001272701,0.00055487565,0.0009882494],"category_scores_gemma":[0.048876952,0.00081272854,0.0008831633,0.010359946,0.0009605597,0.0034675042,0.0023221956,0.001148337,0.0006134511],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060703786,0.00020630566,0.7013665,0.0005448109,0.00038699052,0.0008651254,0.009100785,0.009694615,0.011562203,0.0043793656,0.004284495,0.25700182],"study_design_scores_gemma":[0.000033797176,0.00021283892,0.9303586,0.00023424946,0.00019474432,0.0019543876,0.0038917414,0.020899363,0.011807594,0.00400035,0.02630734,0.000105043866],"about_ca_topic_score_codex":0.0048379656,"about_ca_topic_score_gemma":0.007177921,"teacher_disagreement_score":0.012659769,"about_ca_system_score_codex":0.0012634346,"about_ca_system_score_gemma":0.00083863206,"threshold_uncertainty_score":0.03318906},"labels":[],"label_agreement":null},{"id":"W2402289870","doi":"","title":"RELREA - An Analytical Approach for Evaluating Release Readiness.","year":2014,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software release life cycle; Computer science; Bottleneck; Set (abstract data type); Software; Process management; Fuzzy logic; Point (geometry); Service (business); Software engineering; Software development; Software quality; Engineering; Artificial intelligence","score_opus":0.031876764297685746,"score_gpt":0.2956294902249026,"score_spread":0.2637527259272168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402289870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03249553,0.0011285285,0.94710046,0.00031150682,0.000058202324,0.0006782112,0.00068035704,0.0006792932,0.016867973],"genre_scores_gemma":[0.39199024,0.000618013,0.6029707,0.00007613196,0.000042736552,0.00084829796,0.00060546654,0.000099221244,0.002749096],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.989555,0.0042522876,0.0007049216,0.00095384254,0.0042220275,0.00031189926],"domain_scores_gemma":[0.9839473,0.00902844,0.0024159895,0.000981771,0.0033521792,0.00027440384],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012458556,0.0017142849,0.0010618197,0.017744102,0.0009821558,0.003958683,0.0014671624,0.00076736044,0.0023342986],"category_scores_gemma":[0.026677025,0.00050101714,0.0018632802,0.0067492626,0.0011654117,0.0031566115,0.0021590742,0.0012118061,0.000572876],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002718532,0.00059876626,0.040716983,0.0026146886,0.0009781448,0.0006722575,0.004306494,0.1655895,0.015596104,0.18291262,0.005937008,0.5798057],"study_design_scores_gemma":[0.000041221596,0.0011899727,0.036674865,0.0010753338,0.0005357365,0.001535099,0.005633758,0.7727983,0.016410591,0.13026941,0.03343743,0.00039832343],"about_ca_topic_score_codex":0.002830155,"about_ca_topic_score_gemma":0.0042894767,"teacher_disagreement_score":0.017744102,"about_ca_system_score_codex":0.0027861332,"about_ca_system_score_gemma":0.0023439818,"threshold_uncertainty_score":0.06588799},"labels":[],"label_agreement":null},{"id":"W2402365003","doi":"10.1145/2901739.2903505","title":"The emotional side of software developers in JIRA","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":120,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Politeness; Software; Affect (linguistics); Tracking (education); Data science; Software engineering; Human–computer interaction; Programming language; Psychology","score_opus":0.016600113220322125,"score_gpt":0.2487457158679789,"score_spread":0.23214560264765677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402365003","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9920323,0.000219944,0.0025241198,0.0005104387,0.000030694522,0.000011737111,0.00004568405,0.000066740686,0.0045583],"genre_scores_gemma":[0.9980076,0.00008057529,0.0008056115,0.00014393553,0.00002661905,0.000011643084,0.00005059715,0.000029922323,0.0008434818],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9958554,0.0024701927,0.0001590216,0.00037308686,0.0008878263,0.00025437473],"domain_scores_gemma":[0.96938634,0.02039955,0.005472422,0.0010806175,0.0021097641,0.0015512389],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.00457223,0.00027415706,0.00027195437,0.0010261295,0.00084864546,0.0025246777,0.00026338265,0.00068668695,0.001484868],"category_scores_gemma":[0.028228534,0.00031690756,0.00020361326,0.000523896,0.0012487521,0.0018576168,0.0018504214,0.0011410074,0.000453338],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012153587,0.00030646916,0.62873954,0.0003466675,0.00015918011,0.001002388,0.19376807,0.00074711384,0.032944318,0.0022319753,0.0047298144,0.13380904],"study_design_scores_gemma":[0.000026645308,0.00041341878,0.91748756,0.00012295101,0.00009498707,0.00063649454,0.05492668,0.0055593792,0.0040986817,0.0038154451,0.012701509,0.00011628141],"about_ca_topic_score_codex":0.00068383274,"about_ca_topic_score_gemma":0.0011170891,"teacher_disagreement_score":0.99915135,"about_ca_system_score_codex":0.000526062,"about_ca_system_score_gemma":0.00024017572,"threshold_uncertainty_score":0.024180532},"labels":[],"label_agreement":null},{"id":"W2403963282","doi":"","title":"Étude de la changeabilité des systèmes orientés objet.","year":2008,"lang":"fr","type":"article","venue":"LMO","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Humanities; Political science; Art","score_opus":0.03819496277271119,"score_gpt":0.2855550730078864,"score_spread":0.2473601102351752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2403963282","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8342879,0.003453895,0.14062706,0.0005435681,0.00014675561,0.00034387896,0.001482423,0.0048155375,0.014298908],"genre_scores_gemma":[0.9424142,0.00062156934,0.047678877,0.00008513531,0.000026427893,0.00015940842,0.0015815756,0.0002902704,0.0071425275],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99711096,0.00039488092,0.00016200039,0.00065312843,0.0015494102,0.00012962836],"domain_scores_gemma":[0.98624516,0.0069795586,0.0015526535,0.0016505104,0.003287309,0.00028481716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028353864,0.00074549345,0.00062136643,0.0021301534,0.00066527736,0.0020165471,0.00082550605,0.00077523023,0.00335352],"category_scores_gemma":[0.019609949,0.0004954481,0.00086992135,0.0018990799,0.00057030073,0.0020950164,0.00069653627,0.0007690515,0.0010982209],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010485777,0.00041890077,0.14776398,0.0016388511,0.0006525972,0.0008807152,0.003069879,0.118115634,0.10772881,0.007261887,0.0028641925,0.608556],"study_design_scores_gemma":[0.00011290212,0.0019553558,0.24579266,0.00026285852,0.0008256611,0.0014231682,0.0013715352,0.5097384,0.17444047,0.008110268,0.055751957,0.00021480542],"about_ca_topic_score_codex":0.009560266,"about_ca_topic_score_gemma":0.008508233,"teacher_disagreement_score":0.009560266,"about_ca_system_score_codex":0.0015610824,"about_ca_system_score_gemma":0.0010180459,"threshold_uncertainty_score":0.019009233},"labels":[],"label_agreement":null},{"id":"W2404151716","doi":"10.1145/2889160.2889268","title":"Realistic bug triaging","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Alberta Innovates; Alberta Innovates - Technology Futures","keywords":"Computer science; Task (project management); Ranking (information retrieval); Software bug; Matching (statistics); Data science; World Wide Web; Software engineering; Information retrieval; Software; Programming language; Engineering; Systems engineering","score_opus":0.02396400634740921,"score_gpt":0.2739491754108808,"score_spread":0.24998516906347162,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2404151716","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40616032,0.0026954769,0.53802687,0.004678721,0.00073493517,0.0010048228,0.002921181,0.013385313,0.030392334],"genre_scores_gemma":[0.8477953,0.0004059635,0.14172742,0.0005723177,0.000111622874,0.0003253367,0.0027982178,0.0004904325,0.005773466],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99406177,0.0026158942,0.00034833234,0.0013484547,0.0012639147,0.0003616212],"domain_scores_gemma":[0.9737207,0.01617601,0.0020590688,0.005163752,0.0021350223,0.0007454089],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052767266,0.0009945859,0.00090809166,0.0012151418,0.001089304,0.001818118,0.0024822224,0.0028128764,0.009266813],"category_scores_gemma":[0.05457778,0.00079908204,0.0006353093,0.001293344,0.001040555,0.0026533725,0.0021109781,0.0016268796,0.0024801688],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029341457,0.002126662,0.028100234,0.0020074744,0.00033157034,0.0023036525,0.0028739127,0.38186535,0.0347079,0.03828334,0.09760093,0.40686482],"study_design_scores_gemma":[0.0006524774,0.0017612271,0.011839683,0.00022836954,0.0002281956,0.0027069969,0.000829708,0.83950007,0.009462478,0.08313795,0.049498152,0.00015459595],"about_ca_topic_score_codex":0.003280796,"about_ca_topic_score_gemma":0.0045789713,"teacher_disagreement_score":0.009266813,"about_ca_system_score_codex":0.0012682786,"about_ca_system_score_gemma":0.0013985983,"threshold_uncertainty_score":0.031000614},"labels":[],"label_agreement":null},{"id":"W2404183746","doi":"","title":"Towards a Unified Metrics Suite for JUnit Test Cases.","year":2014,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; Université du Québec à Trois-Rivières","funders":"","keywords":"Unit testing; Test suite; Computer science; Java; Test case; Variance (accounting); Test (biology); Source code; Software; Data mining; Programming language; Machine learning; Regression analysis","score_opus":0.018796076614173172,"score_gpt":0.2512707059323991,"score_spread":0.2324746293182259,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2404183746","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0111602275,0.00044197682,0.98175365,0.00038186065,0.00006486507,0.000670093,0.0005413622,0.0037657362,0.0012201555],"genre_scores_gemma":[0.075919725,0.00036151678,0.91615844,0.00013951171,0.00006819045,0.0013115167,0.0044506146,0.0010955243,0.00049500365],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9381651,0.024790036,0.008651955,0.0030339363,0.024386827,0.0009721376],"domain_scores_gemma":[0.8837446,0.03386369,0.016178614,0.016424797,0.04725709,0.0025312714],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.041582417,0.004330332,0.0026711521,0.017680999,0.0011233006,0.007910326,0.003788548,0.0015997239,0.0010705604],"category_scores_gemma":[0.13541754,0.0013063363,0.0026542938,0.008014465,0.0014698887,0.0068835034,0.005083933,0.004109225,0.0011580923],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026926558,0.0009760756,0.039067447,0.002096152,0.0010386927,0.00055799243,0.0017615901,0.13855158,0.019251544,0.058020514,0.012751293,0.7256578],"study_design_scores_gemma":[0.0001286008,0.0015654307,0.025539484,0.0018552428,0.0005338749,0.0012956209,0.0012098782,0.8006984,0.026470635,0.09404103,0.046320803,0.00034100088],"about_ca_topic_score_codex":0.0026116313,"about_ca_topic_score_gemma":0.0027525693,"teacher_disagreement_score":0.041582417,"about_ca_system_score_codex":0.0027550927,"about_ca_system_score_gemma":0.0051481444,"threshold_uncertainty_score":0.21991152},"labels":[],"label_agreement":null},{"id":"W2404555615","doi":"10.1145/2901739.2901758","title":"An empirical study on the practice of maintaining object-relational mapping code in Java systems","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Blackberry (Canada); Concordia University; Queen's University","funders":"","keywords":"Computer science; Java; Programming language; Object-oriented programming; Empirical research; Code (set theory); Software engineering; Mathematics","score_opus":0.06386517878361975,"score_gpt":0.3532032064081304,"score_spread":0.28933802762451066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2404555615","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9893085,0.00046574982,0.0018208923,0.0010789577,0.00002447255,0.00015203145,0.00008298823,0.00006963847,0.0069966456],"genre_scores_gemma":[0.9963915,0.00022769807,0.0016321735,0.0002687905,0.000017305689,0.00008336275,0.000111826506,0.00006860158,0.001198762],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.96221083,0.0190064,0.002976943,0.0030723407,0.011228054,0.0015053983],"domain_scores_gemma":[0.40870947,0.40439984,0.09569645,0.030984825,0.051319607,0.008889851],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029776057,0.0003662535,0.0003702999,0.0035347051,0.002998627,0.0053373217,0.0027409743,0.0019761869,0.0036249286],"category_scores_gemma":[0.32212138,0.0007543241,0.00038614476,0.0039673187,0.005438065,0.00993745,0.0032747565,0.00330671,0.0009569141],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085375155,0.0034149948,0.71851254,0.00079344766,0.00012140384,0.0007568217,0.14714888,0.0012078307,0.0023939267,0.005936655,0.003357449,0.115502335],"study_design_scores_gemma":[0.00015892833,0.0025008053,0.78982645,0.00088494614,0.00013471529,0.0013284818,0.1606079,0.011254048,0.0022172246,0.0038886734,0.027016155,0.00018168698],"about_ca_topic_score_codex":0.011150174,"about_ca_topic_score_gemma":0.014108213,"teacher_disagreement_score":0.029776057,"about_ca_system_score_codex":0.004922291,"about_ca_system_score_gemma":0.004195754,"threshold_uncertainty_score":0.15747267},"labels":[],"label_agreement":null},{"id":"W2405115200","doi":"10.1109/saner.2016.71","title":"BUMPER: A Tool for Coping with Natural Language Searches of Millions of Bugs and Fixes","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Variety (cybernetics); Reuse; World Wide Web; Software; Software bug; Point (geometry); Open source; Data science; Software engineering; Engineering","score_opus":0.01328571486509194,"score_gpt":0.26611497209477836,"score_spread":0.2528292572296864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405115200","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014455552,0.0040448015,0.03794954,0.0012613266,0.00038007068,0.000818822,0.7046615,0.2317902,0.0046382966],"genre_scores_gemma":[0.012376789,0.0009889196,0.11444204,0.00053949736,0.0000984671,0.0010946854,0.86436397,0.0046584946,0.0014371511],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9932172,0.0011995306,0.0014062317,0.0019068056,0.001934679,0.00033556882],"domain_scores_gemma":[0.9754706,0.0152616035,0.003054733,0.0032629918,0.0020425806,0.0009075573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062332023,0.004678846,0.0020268664,0.027283387,0.0017237008,0.0033162455,0.0054219835,0.0033195193,0.011893037],"category_scores_gemma":[0.033029217,0.001998672,0.0027162228,0.018213604,0.00096173096,0.00900568,0.0065862234,0.003640476,0.012397812],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005713136,0.00030319442,0.009744506,0.005851356,0.000522461,0.0008458563,0.0010826914,0.0026480812,0.0032886532,0.0039058805,0.86353743,0.107698545],"study_design_scores_gemma":[0.0012148574,0.000433275,0.021277465,0.0013965605,0.00039584137,0.0020293158,0.0013922519,0.077383034,0.013113386,0.02537387,0.8556252,0.00036489405],"about_ca_topic_score_codex":0.011343192,"about_ca_topic_score_gemma":0.024479074,"teacher_disagreement_score":0.027283387,"about_ca_system_score_codex":0.0016791663,"about_ca_system_score_gemma":0.0033834164,"threshold_uncertainty_score":0.03978616},"labels":[],"label_agreement":null},{"id":"W2405200576","doi":"10.1145/2896839.2896845","title":"Applying data analytics towards optimized issue management","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates","keywords":"Computer science; Analytics; Data science; Data analysis; Domain (mathematical analysis); Predictive analytics; Process (computing); Set (abstract data type); Data mining","score_opus":0.06569464976977994,"score_gpt":0.3166018727236854,"score_spread":0.25090722295390544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405200576","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16056621,0.0023381906,0.7906845,0.015658362,0.00043708424,0.0024843689,0.0062042237,0.009011988,0.012615041],"genre_scores_gemma":[0.31358176,0.00093882566,0.67714757,0.000587007,0.00021715034,0.0006532247,0.0053084213,0.00050162704,0.0010642854],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97967803,0.0087171765,0.0024969815,0.002784341,0.005663877,0.00065942755],"domain_scores_gemma":[0.9063137,0.05802031,0.009073388,0.012848751,0.011864423,0.0018793255],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019421877,0.0022153608,0.0014835721,0.012587188,0.0017205414,0.011286749,0.0027658776,0.0013391937,0.001663221],"category_scores_gemma":[0.08662004,0.0010243625,0.0016875892,0.010177595,0.0016136671,0.009291519,0.0043165423,0.0039930083,0.0010694767],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004266174,0.0013214945,0.104088776,0.0021636821,0.00052421965,0.0009796682,0.010513335,0.055767268,0.010807624,0.0342574,0.022098728,0.75705117],"study_design_scores_gemma":[0.00019005973,0.00065697514,0.036318544,0.001493772,0.00035528935,0.00080075994,0.02178091,0.5959402,0.036712013,0.19187446,0.11347689,0.0004001496],"about_ca_topic_score_codex":0.0043277564,"about_ca_topic_score_gemma":0.0044325655,"teacher_disagreement_score":0.019421877,"about_ca_system_score_codex":0.0024590937,"about_ca_system_score_gemma":0.005210732,"threshold_uncertainty_score":0.10271394},"labels":[],"label_agreement":null},{"id":"W2405556033","doi":"10.1109/saner.2016.110","title":"Pattern Analysis of TXL Programs","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Programmer; Program comprehension; Programming language; Feature (linguistics); Task (project management); Identification (biology); Natural language processing; Source code; Language identification; Artificial intelligence; Software engineering; Natural language; Software; Software system; Linguistics; Engineering","score_opus":0.02108433717708769,"score_gpt":0.2704798483802654,"score_spread":0.2493955112031777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405556033","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80828345,0.00024426228,0.1793388,0.0003506883,0.000022433916,0.00025895127,0.0021637294,0.0042314054,0.0051062503],"genre_scores_gemma":[0.8748962,0.00012657755,0.119149454,0.00005171387,0.0000103156735,0.00024318504,0.0021605112,0.00035876004,0.0030033195],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985644,0.00025540142,0.00013273461,0.00034489436,0.00058129674,0.00012128264],"domain_scores_gemma":[0.99185383,0.003052985,0.0016357339,0.0012252729,0.0020706053,0.00016156529],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008319193,0.0002791827,0.00026660418,0.0029851724,0.00045365855,0.0009829791,0.0004927095,0.00036288347,0.0015091958],"category_scores_gemma":[0.0072401357,0.00017600416,0.00036634257,0.0020621743,0.00050479115,0.000846978,0.0006040302,0.0003438575,0.00033624057],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056592026,0.00025378188,0.27211395,0.0010398687,0.00010007938,0.002913729,0.008079271,0.012275736,0.107101515,0.008354644,0.004981364,0.58222014],"study_design_scores_gemma":[0.00009931537,0.0007535981,0.3883435,0.00023133903,0.00016003745,0.0066552735,0.004638749,0.3856126,0.16445866,0.014473617,0.03444974,0.00012353623],"about_ca_topic_score_codex":0.0023416793,"about_ca_topic_score_gemma":0.0028401115,"teacher_disagreement_score":0.0029851724,"about_ca_system_score_codex":0.0006041762,"about_ca_system_score_gemma":0.0006702013,"threshold_uncertainty_score":0.005048752},"labels":[],"label_agreement":null},{"id":"W2406365535","doi":"10.1109/saner.2016.80","title":"RACK: Automatic API Recommendation Using Crowdsourced Knowledge","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":151,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Java; Code (set theory); Information retrieval; Matching (statistics); Precision and recall; Programming language; Search engine; Database; Set (abstract data type)","score_opus":0.055700943670632184,"score_gpt":0.33794730808844026,"score_spread":0.2822463644178081,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2406365535","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09063884,0.004419764,0.73590666,0.0021691206,0.0006449706,0.0030226642,0.037607476,0.10622494,0.019365622],"genre_scores_gemma":[0.2873628,0.00090416207,0.66330624,0.00077442394,0.00022514016,0.0012498229,0.03625303,0.0015493718,0.008374909],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9944883,0.0011334835,0.000384665,0.0019433704,0.0017555458,0.00029469922],"domain_scores_gemma":[0.9894119,0.005384484,0.00076538505,0.0023048045,0.0017561943,0.00037720776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034683628,0.0027061237,0.0021041099,0.010929969,0.0014869863,0.0017603826,0.003523005,0.0023381761,0.0053325826],"category_scores_gemma":[0.017098319,0.00079864357,0.001722613,0.0067356005,0.00076840713,0.0036175805,0.0030764795,0.0015444916,0.0061211446],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013851044,0.0011753693,0.015532641,0.0025920516,0.00069594145,0.00095693907,0.0014015078,0.031889524,0.02661422,0.004427087,0.14801645,0.7653131],"study_design_scores_gemma":[0.0004625983,0.00036265462,0.009236297,0.00030142107,0.00032354388,0.00062658434,0.0016397102,0.86985964,0.025399767,0.024171632,0.06727451,0.00034173598],"about_ca_topic_score_codex":0.029839013,"about_ca_topic_score_gemma":0.048038304,"teacher_disagreement_score":0.029839013,"about_ca_system_score_codex":0.0014261649,"about_ca_system_score_gemma":0.003888632,"threshold_uncertainty_score":0.059330642},"labels":[],"label_agreement":null},{"id":"W2407266979","doi":"10.1145/2897134.2897136","title":"The role of semiotic engineering in software engineering","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Semiotics; Computer science; Software engineering; Linguistics; Philosophy","score_opus":0.005392298951782116,"score_gpt":0.20395107991227618,"score_spread":0.19855878096049406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2407266979","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016838297,0.014089483,0.843157,0.022279475,0.00093714386,0.00016085753,0.00007083846,0.00026630922,0.10220059],"genre_scores_gemma":[0.66631114,0.011337353,0.31000715,0.0026854728,0.0010340376,0.0005901021,0.00011121108,0.0002531693,0.0076703634],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9678093,0.024498483,0.0016863694,0.0017548333,0.003684838,0.00056616165],"domain_scores_gemma":[0.931752,0.054136205,0.002533995,0.006180085,0.0041308347,0.0012668706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022461452,0.0009936937,0.0013037508,0.004830925,0.004338166,0.013666332,0.0018144465,0.003744225,0.0019675817],"category_scores_gemma":[0.03290251,0.0009305913,0.0011310547,0.0031498622,0.06448758,0.016215913,0.00742381,0.0063296235,0.0007496222],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008317189,0.000010101472,0.00016801048,0.00009108292,0.000007719856,0.000071438575,0.002935396,0.0010354992,0.00021352305,0.98904675,0.00039742995,0.0060147075],"study_design_scores_gemma":[0.000010804318,0.000022228554,0.000102578444,0.00014670538,0.000006507588,0.00015337906,0.0010108595,0.002938818,0.00041217965,0.9735386,0.021630166,0.000027124752],"about_ca_topic_score_codex":0.0019841902,"about_ca_topic_score_gemma":0.0009532655,"teacher_disagreement_score":0.022461452,"about_ca_system_score_codex":0.0049232487,"about_ca_system_score_gemma":0.0046016737,"threshold_uncertainty_score":0.11878896},"labels":[],"label_agreement":null},{"id":"W2407299292","doi":"10.1145/2901739.2903499","title":"Analysis of exception handling patterns in Java projects","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Java; Computer science; Exception handling; Programming language; Class (philosophy); Class hierarchy; Software bug; Generics in Java; Hierarchy; Software; Scala; Empirical research; Real time Java; Software engineering; Java annotation; Object-oriented programming; Artificial intelligence; Mathematics; Statistics","score_opus":0.026902615553880547,"score_gpt":0.2799268856200367,"score_spread":0.25302427006615613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2407299292","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9956151,0.00017452285,0.002524998,0.00007637118,0.0000042621937,0.000062264015,0.0006815011,0.000040537416,0.0008203642],"genre_scores_gemma":[0.99122417,0.00017505414,0.006082363,0.000025503758,0.0000066877215,0.00013502558,0.0018025459,0.000028750028,0.00051992026],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9939521,0.001270791,0.0012292311,0.0012556128,0.0018785644,0.0004137211],"domain_scores_gemma":[0.93619275,0.03405575,0.018039191,0.0026769037,0.0077547682,0.0012807786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030269227,0.00022709143,0.00031004837,0.007909794,0.0006772084,0.0014651874,0.0005919052,0.0005127019,0.0006335833],"category_scores_gemma":[0.029078383,0.00028788534,0.0003595484,0.007577331,0.000680376,0.0014941684,0.0012313436,0.00051229744,0.00025394082],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009117638,0.00012370765,0.9564175,0.00019557838,0.000038060705,0.00035782118,0.0059776744,0.0003421236,0.0022661153,0.00053475145,0.00041172202,0.03324384],"study_design_scores_gemma":[0.000007124644,0.00011818822,0.9822953,0.000116017814,0.0000316632,0.0009525803,0.008597072,0.0033104613,0.0011495561,0.0007653621,0.0026292922,0.000027496768],"about_ca_topic_score_codex":0.0034813625,"about_ca_topic_score_gemma":0.0052340776,"teacher_disagreement_score":0.007909794,"about_ca_system_score_codex":0.0006112692,"about_ca_system_score_gemma":0.001261503,"threshold_uncertainty_score":0.016008079},"labels":[],"label_agreement":null},{"id":"W2408052243","doi":"10.1145/2889160.2889243","title":"A study of the quality-impacting practices of modern code review at Sony mobile","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; McGill University","funders":"Japan Society for the Promotion of Science","keywords":"Software quality; Computer science; Context (archaeology); Code review; Flexibility (engineering); Quality (philosophy); Software engineering; Code (set theory); Process (computing); Software; World Wide Web; Software development; Operating system","score_opus":0.1167875131363045,"score_gpt":0.4270779480597355,"score_spread":0.310290434923431,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408052243","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99633765,0.00023376125,0.0015776786,0.00039837006,0.000008438603,0.00008261913,0.000017522649,0.000033784167,0.0013101762],"genre_scores_gemma":[0.9957807,0.00022534492,0.0029071236,0.00014988663,0.000020787867,0.000103675025,0.000032452124,0.000028293212,0.0007517535],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.95697093,0.025970083,0.001936883,0.0029013872,0.010866084,0.0013546571],"domain_scores_gemma":[0.56850344,0.281466,0.085208364,0.013385847,0.04033625,0.01110007],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0335448,0.00031933328,0.00040290577,0.004312337,0.002350048,0.0027520328,0.001219162,0.0009239136,0.0009873275],"category_scores_gemma":[0.14943141,0.0006556491,0.0003642963,0.0025386214,0.0033641832,0.0030872063,0.0031880066,0.0011255774,0.00023219902],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000401994,0.0009343384,0.43529725,0.00081136863,0.00012718621,0.0017223636,0.34727365,0.000745498,0.012136621,0.0014024349,0.0015237599,0.19762352],"study_design_scores_gemma":[0.00008859514,0.0039662686,0.8577197,0.0006200151,0.000100105965,0.001725487,0.10902752,0.0037573155,0.005243335,0.0007958156,0.016804969,0.00015081957],"about_ca_topic_score_codex":0.0056727184,"about_ca_topic_score_gemma":0.00905894,"teacher_disagreement_score":0.9664552,"about_ca_system_score_codex":0.005091947,"about_ca_system_score_gemma":0.006156748,"threshold_uncertainty_score":0.17740393},"labels":[],"label_agreement":null},{"id":"W2408185736","doi":"","title":"The Impact of Confirmation Bias on the Release-based Defect Prediction of Developer Groups.","year":2013,"lang":"en","type":"article","venue":"ENLIGHTEN (Jurnal Bimbingan dan Konseling Islam)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science","score_opus":0.03059326370806244,"score_gpt":0.25242433474541287,"score_spread":0.22183107103735042,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408185736","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98502445,0.000625631,0.01198639,0.00042698559,0.000036939997,0.000055569068,0.0004494369,0.00013355725,0.0012609777],"genre_scores_gemma":[0.99707913,0.00006263271,0.0020703091,0.00004214797,0.00001624176,0.000020227159,0.00041147968,0.000017345646,0.00028043034],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9849268,0.009249786,0.0007674637,0.0022102282,0.0022974222,0.00054829405],"domain_scores_gemma":[0.5570472,0.36813846,0.042086255,0.01883373,0.0112873465,0.002607043],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031875532,0.0007981981,0.0006477627,0.0022319367,0.00059871306,0.0015684719,0.0012935904,0.001281491,0.0012981533],"category_scores_gemma":[0.20079331,0.00039924603,0.00086205936,0.0017738002,0.001169787,0.0021592653,0.0014739254,0.0018288498,0.00056228985],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004260304,0.00013034286,0.9717942,0.000055772765,0.0002647623,0.00008305491,0.00049329887,0.0041732853,0.00024982274,0.00053459476,0.00068142475,0.021113506],"study_design_scores_gemma":[0.000075147436,0.0005072097,0.849754,0.0000848484,0.00035783343,0.0002902423,0.0007001356,0.14143594,0.0017573582,0.003716648,0.0012553008,0.00006534329],"about_ca_topic_score_codex":0.009715957,"about_ca_topic_score_gemma":0.0091905715,"teacher_disagreement_score":0.031875532,"about_ca_system_score_codex":0.0008931308,"about_ca_system_score_gemma":0.0009672066,"threshold_uncertainty_score":0.16857594},"labels":[],"label_agreement":null},{"id":"W2408538694","doi":"10.1109/saner.2016.73","title":"On the Detection of Licenses Violations in the Android Ecosystem","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Android (operating system); License; Reuse; Mobile apps; Computer science; Open source; App store; Order (exchange); World Wide Web; Internet privacy; Computer security; Code reuse; Business; Operating system; Software; Engineering","score_opus":0.017674177407946832,"score_gpt":0.24154054337155445,"score_spread":0.22386636596360762,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408538694","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9947507,0.0005255506,0.0014670779,0.00024883813,0.000019992427,0.000060856724,0.00059598713,0.00015052778,0.0021804215],"genre_scores_gemma":[0.9950317,0.00026345314,0.0026896794,0.00006514694,0.000015911593,0.000042939184,0.0011163253,0.000059840437,0.0007150364],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98984027,0.0017425689,0.0014178257,0.0010681157,0.0050762957,0.00085495436],"domain_scores_gemma":[0.8751852,0.050372597,0.05162296,0.005920161,0.0145112015,0.002387904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005332375,0.00043852188,0.00044286918,0.008436751,0.0011157995,0.002034206,0.0008505069,0.001080308,0.00080132234],"category_scores_gemma":[0.068447046,0.00055256765,0.00057716295,0.005423943,0.0009872759,0.0036017662,0.0022546733,0.0016179002,0.0006522357],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011419349,0.0001232143,0.96447134,0.00011842851,0.00004439066,0.001031404,0.0030020303,0.0006835636,0.0018851121,0.0003677149,0.0009544673,0.027204094],"study_design_scores_gemma":[0.0000063330085,0.0001125332,0.9793396,0.00009386087,0.00004438035,0.0016317242,0.0032312255,0.010249832,0.0017181182,0.00042573662,0.0030918946,0.00005480561],"about_ca_topic_score_codex":0.019683648,"about_ca_topic_score_gemma":0.018889017,"teacher_disagreement_score":0.019683648,"about_ca_system_score_codex":0.0010896304,"about_ca_system_score_gemma":0.0014759803,"threshold_uncertainty_score":0.039138198},"labels":[],"label_agreement":null},{"id":"W2408619423","doi":"10.1109/saner.2016.18","title":"An Empirical Study on the Use of CSS Preprocessors","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Empirical research; Mathematics; Statistics","score_opus":0.13149249201718202,"score_gpt":0.3674709244662258,"score_spread":0.2359784324490438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408619423","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99772125,0.00019034714,0.00042575662,0.00007480578,0.000004966783,0.000038499336,0.00012894784,0.000014343725,0.0014009321],"genre_scores_gemma":[0.99628395,0.0004346315,0.0016469839,0.000097607226,0.000008175469,0.00008937531,0.00042801662,0.000033484943,0.0009777652],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99135053,0.0038281428,0.00091626745,0.0008695472,0.0026156902,0.00041992712],"domain_scores_gemma":[0.82346535,0.130872,0.020467214,0.005706556,0.016976709,0.0025121025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009752839,0.00035648007,0.0003842349,0.0027077263,0.000943239,0.0026601686,0.00102584,0.0009801572,0.0018905902],"category_scores_gemma":[0.07900062,0.000474638,0.00037073012,0.0030286901,0.0012736571,0.0039546033,0.0010848205,0.0014088571,0.0006900859],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003549334,0.00091656094,0.88797534,0.0007254316,0.00011947405,0.0008046522,0.05228046,0.0002833932,0.0023939596,0.0007575032,0.0017989088,0.051589414],"study_design_scores_gemma":[0.000024366438,0.0007251206,0.9062415,0.00039376778,0.000111524874,0.0013197458,0.0728398,0.002840714,0.0024025938,0.00032902916,0.012694335,0.0000776173],"about_ca_topic_score_codex":0.0030231231,"about_ca_topic_score_gemma":0.0044582267,"teacher_disagreement_score":0.009752839,"about_ca_system_score_codex":0.00077772565,"about_ca_system_score_gemma":0.0008022709,"threshold_uncertainty_score":0.05157858},"labels":[],"label_agreement":null},{"id":"W2411173501","doi":"10.4018/978-1-60566-904-5.ch013","title":"Dimensions of UML Diagram Use","year":2010,"lang":"en","type":"book-chapter","venue":"Advances in database research (ADR) book series/Advances in database research series","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland; University of Lethbridge","funders":"","keywords":"Unified Modeling Language; Applications of UML; UML tool; Class diagram; Computer science; Communication diagram; Systems Modeling Language; Use Case Diagram; Software engineering; Object Constraint Language; Programming language; Software","score_opus":0.06451278341028834,"score_gpt":0.38779775375632203,"score_spread":0.3232849703460337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2411173501","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08148675,0.01466839,0.13788207,0.013379592,0.00056495564,0.00043827298,0.0012954244,0.0013317844,0.7489528],"genre_scores_gemma":[0.7141721,0.014574615,0.20671727,0.0021569717,0.00047662598,0.00096253003,0.0026092264,0.0009614763,0.05736919],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9757681,0.010875201,0.0015726559,0.001097851,0.010215681,0.0004706032],"domain_scores_gemma":[0.97246325,0.017138619,0.0026057244,0.00231786,0.00466055,0.0008139453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008872441,0.00054243463,0.000252637,0.0061125327,0.0010646533,0.009035043,0.00078280334,0.0011521708,0.0050724605],"category_scores_gemma":[0.027107015,0.0004204277,0.00037866406,0.007061267,0.0027291402,0.008837,0.0037411437,0.0017078886,0.0015244372],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004989035,0.00007463689,0.011189374,0.0006022676,0.00003721823,0.00022472271,0.046425942,0.00088567345,0.0031755092,0.5244215,0.024724923,0.38818833],"study_design_scores_gemma":[0.000012989219,0.00004892678,0.009929724,0.0011774409,0.000021514596,0.0016107758,0.012261612,0.001993492,0.0011996692,0.13415502,0.8375254,0.00006335226],"about_ca_topic_score_codex":0.0018040832,"about_ca_topic_score_gemma":0.0016971827,"teacher_disagreement_score":0.009035043,"about_ca_system_score_codex":0.0024316483,"about_ca_system_score_gemma":0.0018965112,"threshold_uncertainty_score":0.046922505},"labels":[],"label_agreement":null},{"id":"W2417033663","doi":"10.1007/s10664-016-9438-4","title":"License usage and changes: a large-scale study on gitHub","year":2016,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"European Commission; National Science Foundation","keywords":"License; Computer science; Traceability; Commit; Java; Software engineering; Python (programming language); AspectJ; Software; JavaScript; Secure coding; World Wide Web; Empirical research; Software development; Reuse; Computer security; Database; Programming language; Engineering; Software security assurance; Operating system","score_opus":0.027664631968529942,"score_gpt":0.28769272460390743,"score_spread":0.2600280926353775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2417033663","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99818265,0.00007532822,0.00007192691,0.0000690788,0.000002764768,0.000022161988,0.000637134,0.00001661874,0.00092243176],"genre_scores_gemma":[0.9958152,0.00015587873,0.00018499694,0.0001394382,0.000012062509,0.000046226127,0.0024054411,0.000061438324,0.0011792447],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99843735,0.00040886173,0.00009333655,0.00025690565,0.00050596474,0.00029761696],"domain_scores_gemma":[0.9854894,0.0059023383,0.004093383,0.0009999565,0.001755936,0.0017589796],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0013079765,0.00042632245,0.00052962435,0.004043209,0.0014341117,0.0018976216,0.0012248083,0.0009143418,0.0028441984],"category_scores_gemma":[0.0083797835,0.00036175986,0.00048812857,0.0074919616,0.0015587511,0.0026304852,0.0020119988,0.0013577421,0.0012729028],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028715222,0.0013470894,0.96220756,0.0001470661,0.00015534875,0.0010916993,0.0143310595,0.00029923904,0.00081454037,0.00040753867,0.0036058647,0.015305981],"study_design_scores_gemma":[0.00000970698,0.00007769039,0.98693323,0.000030751344,0.00002924388,0.00019546064,0.009949966,0.0005132527,0.00021065355,0.0000575207,0.0019726383,0.000019826122],"about_ca_topic_score_codex":0.08933243,"about_ca_topic_score_gemma":0.12779553,"teacher_disagreement_score":0.9959568,"about_ca_system_score_codex":0.0019591174,"about_ca_system_score_gemma":0.0015804162,"threshold_uncertainty_score":0.17762494},"labels":[],"label_agreement":null},{"id":"W2427333829","doi":"10.1109/saner.2016.78","title":"Bug Replication in Code Clones: An Empirical Study","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Code refactoring; Cloning (programming); Programming language; Computer science; Replication (statistics); Code (set theory); Java; Software bug; Biology; Software; Genetics; Virology; Gene","score_opus":0.058476840466890764,"score_gpt":0.3802583650532162,"score_spread":0.3217815245863254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2427333829","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9978269,0.0002978642,0.0010317756,0.000079947116,0.0000039578217,0.00006875998,0.00012858128,0.000019678262,0.0005425516],"genre_scores_gemma":[0.9979948,0.00021240114,0.0011568171,0.00004031363,0.000009260408,0.00006583608,0.00024569745,0.000016150021,0.00025879592],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98266554,0.007651488,0.0018917839,0.0023920925,0.004807364,0.00059180916],"domain_scores_gemma":[0.6522311,0.2637864,0.048754856,0.01202367,0.020605592,0.002598303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011609847,0.00045752694,0.00046062193,0.0036205878,0.0012039182,0.0016210976,0.0012215845,0.0011856734,0.0012895603],"category_scores_gemma":[0.112498954,0.0005048376,0.00050218374,0.003416972,0.0022948748,0.0027811155,0.0017892946,0.0016657265,0.00028473884],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019424019,0.0005897838,0.95045334,0.0003731473,0.00011868977,0.0011312822,0.017759606,0.000660194,0.00094810623,0.00050241634,0.00085562654,0.026413403],"study_design_scores_gemma":[0.00005556417,0.0009970106,0.95453054,0.00029900746,0.00017869593,0.005007801,0.022102999,0.00845347,0.0019119997,0.0007582068,0.0056284326,0.0000761614],"about_ca_topic_score_codex":0.003190467,"about_ca_topic_score_gemma":0.004135305,"teacher_disagreement_score":0.011609847,"about_ca_system_score_codex":0.001064053,"about_ca_system_score_gemma":0.0011321651,"threshold_uncertainty_score":0.06139952},"labels":[],"label_agreement":null},{"id":"W2432194708","doi":"10.11575/prism/3210","title":"Design and evaluation of explanation-based decision support for software release planning","year":2009,"lang":"en","type":"article","venue":"PRISM (University of Calgary)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Decision support system; Software; Black box; Software system; Software engineering; Management science; Artificial intelligence; Engineering","score_opus":0.034731843871939144,"score_gpt":0.27071729945431533,"score_spread":0.23598545558237619,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2432194708","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51956743,0.0004430678,0.46433398,0.0005204049,0.00008076255,0.0059188264,0.00029077358,0.0055280654,0.0033166828],"genre_scores_gemma":[0.46244043,0.00014038167,0.53408957,0.00010509711,0.000013791817,0.0019190793,0.00032567445,0.00010253407,0.0008634305],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943235,0.0035089313,0.00051642134,0.0006300609,0.0007872007,0.00023393768],"domain_scores_gemma":[0.95764005,0.03522471,0.0022325616,0.0015532967,0.0025570414,0.00079232786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008951325,0.0009769106,0.00055630493,0.0009328626,0.00035192544,0.0016006678,0.0025251242,0.0013326133,0.003944056],"category_scores_gemma":[0.03689619,0.00060233,0.00056787237,0.00047608593,0.00068742485,0.0017262273,0.0014019375,0.00094674184,0.0003760574],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010974272,0.005156182,0.014030347,0.004534425,0.00049229426,0.001063302,0.007675856,0.117093764,0.061768226,0.013335758,0.0027170498,0.7611585],"study_design_scores_gemma":[0.0038965459,0.01357104,0.0129927555,0.0005625962,0.000728676,0.00042804878,0.0020470056,0.8816128,0.061021477,0.0081916535,0.014715538,0.00023183845],"about_ca_topic_score_codex":0.0009782508,"about_ca_topic_score_gemma":0.00095615955,"teacher_disagreement_score":0.008951325,"about_ca_system_score_codex":0.0013519472,"about_ca_system_score_gemma":0.0017617004,"threshold_uncertainty_score":0.047339678},"labels":[],"label_agreement":null},{"id":"W2440701844","doi":"10.1109/saner.2016.87","title":"An Empirical Study on Recommendations of Similar Bugs","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Conselho Nacional de Desenvolvimento Científico e Tecnológico","keywords":"Software bug; Computer science; Recommender system; Software engineering; Security bug; Field (mathematics); Code review; Code (set theory); Feature (linguistics); Software maintenance; Open source; Software regression; Empirical research; World Wide Web; Data science; Software; Information retrieval; Software development; Static program analysis; Programming language; Software quality; Computer security","score_opus":0.06122640586803564,"score_gpt":0.3859263476921034,"score_spread":0.32469994182406775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2440701844","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9955622,0.00040096516,0.0009246266,0.00035023302,0.000018015842,0.000145371,0.00085965055,0.000055650155,0.0016833838],"genre_scores_gemma":[0.9943545,0.00021114174,0.0024143932,0.00016376379,0.000020251704,0.00015951834,0.0018036606,0.000028235074,0.00084461406],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97478205,0.01599896,0.0024108659,0.0022883771,0.0039870404,0.00053273304],"domain_scores_gemma":[0.42590854,0.5034452,0.032395635,0.016756997,0.017215144,0.0042785117],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025131145,0.0003860813,0.0007307871,0.0025392685,0.0010340569,0.0019456937,0.0014081146,0.0016969516,0.0036756475],"category_scores_gemma":[0.25488877,0.0004922351,0.0005078156,0.0031881942,0.0010065531,0.0036582404,0.0010239835,0.002016408,0.0009862764],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017451186,0.00297497,0.9373489,0.0005844373,0.00035569415,0.0002778152,0.004734317,0.0011493387,0.0005457679,0.00038085174,0.004184518,0.045718275],"study_design_scores_gemma":[0.00060200185,0.0050105597,0.9412621,0.00031237493,0.00025377524,0.0013237987,0.010824238,0.027405005,0.0013215282,0.0006200648,0.010923173,0.00014135352],"about_ca_topic_score_codex":0.0072174612,"about_ca_topic_score_gemma":0.007424355,"teacher_disagreement_score":0.025131145,"about_ca_system_score_codex":0.0007906111,"about_ca_system_score_gemma":0.0007265727,"threshold_uncertainty_score":0.13290781},"labels":[],"label_agreement":null},{"id":"W2461181047","doi":"10.82308/7763","title":"Using genetic algorithms to optimize software quality estimation models","year":2004,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Data mining; Software; Context (archaeology); Machine learning; Software quality; Set (abstract data type); Genetic algorithm; Quality (philosophy); Software metric; Software development; Artificial intelligence","score_opus":0.05608361537139245,"score_gpt":0.3190503475250635,"score_spread":0.26296673215367106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2461181047","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06944686,0.0004892594,0.9240347,0.00046010566,0.0000519484,0.00014248048,0.000105071136,0.00055264647,0.0047168722],"genre_scores_gemma":[0.5061377,0.0005516593,0.48839357,0.00024824313,0.000049575716,0.00056420855,0.00055175135,0.0002303517,0.0032728962],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99876994,0.00058892835,0.00006493593,0.00022052298,0.00022194417,0.00013372336],"domain_scores_gemma":[0.99448663,0.0044882745,0.00032816725,0.00015470803,0.00046588707,0.00007634939],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028005315,0.001769123,0.0014547858,0.0022678517,0.00058740855,0.0016107775,0.0015762593,0.001958036,0.0014624066],"category_scores_gemma":[0.011646626,0.0010144467,0.001238244,0.0017427756,0.0010225114,0.001190952,0.0013041823,0.0015669023,0.00028915258],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000010413573,0.00002472147,0.00051090325,0.000012945028,0.000022669952,0.000017796463,0.000027085069,0.9851682,0.00015905681,0.002021005,0.00016437775,0.011860728],"study_design_scores_gemma":[0.000005310531,0.000010761237,0.00006654958,0.0000061669534,0.0000075310472,0.0000033181054,0.000008597724,0.9980469,0.000095897216,0.0016283097,0.000118285876,0.000002433335],"about_ca_topic_score_codex":0.016184062,"about_ca_topic_score_gemma":0.011274591,"teacher_disagreement_score":0.016184062,"about_ca_system_score_codex":0.0020543651,"about_ca_system_score_gemma":0.0019904436,"threshold_uncertainty_score":0.032179713},"labels":[],"label_agreement":null},{"id":"W2465706857","doi":"","title":"Understanding open source software peer review: review processes, parameters and statistical models, and underlying behaviours and mechanisms","year":2011,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Construct (python library); Computer science; Quality (philosophy); Code review; Data science; Technical peer review; Knowledge management; Software development; Software; Process management; Peer review; Engineering; Software quality; Political science","score_opus":0.23862595243352208,"score_gpt":0.36821974594049645,"score_spread":0.12959379350697436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2465706857","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4152446,0.006883454,0.542907,0.013933369,0.00017283187,0.0013685402,0.00054983824,0.0010193012,0.017921098],"genre_scores_gemma":[0.9554334,0.0014948148,0.040486354,0.00025678193,0.00013040817,0.00051702274,0.00020448344,0.00010951798,0.0013672585],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.92684764,0.04347261,0.0048501915,0.008459267,0.014286775,0.0020834517],"domain_scores_gemma":[0.34200293,0.5124867,0.07853323,0.021770786,0.041486684,0.0037196418],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09172928,0.0009238815,0.001701154,0.0090056285,0.0019233539,0.012419603,0.0036988123,0.0037219403,0.0022455277],"category_scores_gemma":[0.37425822,0.0016856398,0.0013775522,0.005185089,0.006686264,0.02215386,0.003138256,0.0033906987,0.0009207188],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004409349,0.0008024062,0.25062826,0.0022362517,0.0006885007,0.00068328454,0.029990895,0.14302225,0.005199166,0.33522788,0.006970875,0.22410925],"study_design_scores_gemma":[0.00013100439,0.00046377553,0.116820775,0.000869364,0.00021343818,0.0007941732,0.0097133005,0.49408764,0.0028955291,0.3599889,0.013530426,0.00049172266],"about_ca_topic_score_codex":0.01038905,"about_ca_topic_score_gemma":0.0059346943,"teacher_disagreement_score":0.9082707,"about_ca_system_score_codex":0.0076307976,"about_ca_system_score_gemma":0.00765396,"threshold_uncertainty_score":0.4851166},"labels":[],"label_agreement":null},{"id":"W2465807508","doi":"10.1109/icpc.2016.7503717","title":"A cooperative approach for combining client-based and library-based API usage pattern mining","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Application programming interface; Dependency (UML); Code (set theory); Generalizability theory; Software; Software engineering; Data mining; Programming language","score_opus":0.025912212547362642,"score_gpt":0.25125050310702646,"score_spread":0.2253382905596638,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2465807508","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03613776,0.00024185139,0.95351857,0.00025285548,0.000015450607,0.00038295545,0.00020508483,0.007878093,0.0013674388],"genre_scores_gemma":[0.21529254,0.00012574637,0.7807057,0.00017064197,0.000022833647,0.00042825832,0.00089560886,0.0004404784,0.0019182254],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9907718,0.0022847196,0.00073132984,0.0019780973,0.0037719756,0.00046210628],"domain_scores_gemma":[0.9861224,0.0043861964,0.0012164373,0.004386098,0.003491171,0.00039770128],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005768278,0.0016905554,0.0018855525,0.008505125,0.0011726962,0.0019648795,0.0046591237,0.0018177304,0.0008828936],"category_scores_gemma":[0.015032304,0.0013079721,0.0019333399,0.006704498,0.00090929976,0.004396344,0.003819667,0.0018068589,0.0010987023],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000346161,0.001398915,0.034350816,0.0005389349,0.0006332018,0.0006606762,0.0022731754,0.011845357,0.05874477,0.005088953,0.0056235185,0.87849563],"study_design_scores_gemma":[0.00009931829,0.00057167193,0.017091947,0.00012250718,0.0003961242,0.0024697394,0.001035222,0.8686523,0.0780799,0.0143189905,0.016963128,0.00019921683],"about_ca_topic_score_codex":0.0044035986,"about_ca_topic_score_gemma":0.007110456,"teacher_disagreement_score":0.008505125,"about_ca_system_score_codex":0.0006436672,"about_ca_system_score_gemma":0.002401827,"threshold_uncertainty_score":0.030505955},"labels":[],"label_agreement":null},{"id":"W2466101837","doi":"","title":"Convergent Software Peer Review Practices","year":2013,"lang":"en","type":"article","venue":"Foundations of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Android (operating system); Software; Software peer review; World Wide Web; Data science; Software review; Best practice; Incentive; Software engineering; Knowledge management; Software development; Software construction; Operating system","score_opus":0.03562041747637083,"score_gpt":0.30526978279252304,"score_spread":0.2696493653161522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2466101837","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8361916,0.005609771,0.07112269,0.0064998902,0.0002923764,0.0012451775,0.00018623084,0.0009583915,0.07789391],"genre_scores_gemma":[0.9853881,0.0007026629,0.009273533,0.00036701516,0.00013621247,0.00027219762,0.00007270605,0.00011353043,0.0036739514],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.84742266,0.067466445,0.010842452,0.016410802,0.053539652,0.004317988],"domain_scores_gemma":[0.48511776,0.24692324,0.07597433,0.07933018,0.09755832,0.0150960125],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06646929,0.000561959,0.0009178796,0.0069694268,0.005562403,0.0073734336,0.0036023806,0.0019302013,0.003968915],"category_scores_gemma":[0.3003454,0.0007525643,0.00070648943,0.004465494,0.0057592713,0.0070231524,0.008955006,0.0022365735,0.0014109436],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043426416,0.00047213477,0.27445787,0.002069168,0.0005887766,0.002003159,0.21697962,0.0021436904,0.012089005,0.040154133,0.008943,0.4396651],"study_design_scores_gemma":[0.0002432663,0.0019913025,0.4898387,0.0027264068,0.00045568345,0.0075247497,0.14718589,0.012830156,0.012919552,0.090221554,0.23342004,0.0006427031],"about_ca_topic_score_codex":0.0023930443,"about_ca_topic_score_gemma":0.0029861063,"teacher_disagreement_score":0.9335307,"about_ca_system_score_codex":0.004245125,"about_ca_system_score_gemma":0.006791669,"threshold_uncertainty_score":0.3515274},"labels":[],"label_agreement":null},{"id":"W2468368474","doi":"","title":"Improving design quality using meta-pattern transformations: a metric-based approach: Research Articles","year":2004,"lang":"en","type":"article","venue":"Conference on Software Maintenance and Reengineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Metric (unit); Business process reengineering; Task (project management); Software engineering; Quality (philosophy); Process (computing); Object-oriented design; Software quality; Metamodeling; Software maintenance; Data mining; Software; Software system; Software development; Systems engineering; Programming language; Engineering","score_opus":0.25274240047442287,"score_gpt":0.3497299494202283,"score_spread":0.09698754894580541,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2468368474","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06974101,0.024073146,0.88275796,0.0051216697,0.00026516916,0.0003042752,0.00021555035,0.0022601488,0.015261119],"genre_scores_gemma":[0.39765793,0.013327173,0.5837889,0.00030379187,0.00022174453,0.00022303572,0.00043702687,0.00064775907,0.0033926978],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935875,0.0019531904,0.000492608,0.0006164663,0.0032337883,0.000116408206],"domain_scores_gemma":[0.9840682,0.006255931,0.0019569437,0.0021942223,0.0051645064,0.00036013327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049944464,0.0008314256,0.0008302634,0.0040203375,0.00041726048,0.003289818,0.001229803,0.001112866,0.0013774538],"category_scores_gemma":[0.019352473,0.0003466085,0.0005649148,0.0058586886,0.0015112574,0.004024741,0.0007471706,0.0009760523,0.0004333759],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000096111115,0.00023646836,0.0054717255,0.0008999094,0.000081306396,0.000052401774,0.00028048258,0.013018674,0.00894472,0.034340568,0.0036638584,0.9329137],"study_design_scores_gemma":[0.00021131927,0.001508149,0.031791948,0.00145264,0.00043364198,0.0018856402,0.0010756054,0.5210118,0.07126798,0.26156425,0.10750696,0.0002900839],"about_ca_topic_score_codex":0.0015035928,"about_ca_topic_score_gemma":0.0011627853,"teacher_disagreement_score":0.0049944464,"about_ca_system_score_codex":0.0014116915,"about_ca_system_score_gemma":0.0015394167,"threshold_uncertainty_score":0.0264135},"labels":[],"label_agreement":null},{"id":"W2469414321","doi":"10.1016/j.jss.2016.06.101","title":"Reverse engineering reusable software components from object-oriented APIs","year":2016,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Computer science; Reuse; Component-based software engineering; Software engineering; Component (thermodynamics); Object-oriented programming; Java; Software; Application programming interface; Software development; Programming language; Engineering","score_opus":0.014664627731742352,"score_gpt":0.22200487766130164,"score_spread":0.2073402499295593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2469414321","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07411077,0.00032656055,0.9128574,0.0002037945,0.00016926687,0.00034797416,0.00012746909,0.00565941,0.0061973706],"genre_scores_gemma":[0.28677318,0.000750062,0.6956053,0.00020934621,0.00004473851,0.00022403328,0.0010929778,0.002812557,0.0124878185],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99805486,0.00032158967,0.00014044403,0.00017074433,0.0010800635,0.00023232303],"domain_scores_gemma":[0.99328965,0.0018732651,0.00046558646,0.0029286102,0.0013570664,0.000085830805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022560991,0.00086592947,0.0004951065,0.0014977027,0.0006130356,0.0019781797,0.0014187698,0.001004825,0.0022521922],"category_scores_gemma":[0.008075135,0.00083351345,0.0015596434,0.0009138533,0.0010162329,0.0020526436,0.0021698126,0.0020670325,0.0017398852],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003964919,0.00048515334,0.008056482,0.0012161926,0.00026108156,0.0020112202,0.001488396,0.04051544,0.16549288,0.12364183,0.00712194,0.6493129],"study_design_scores_gemma":[0.00020295334,0.00047665642,0.0031548827,0.000385259,0.00072708534,0.0024158363,0.0007301726,0.34396175,0.46382582,0.10561099,0.07838245,0.00012611662],"about_ca_topic_score_codex":0.002048435,"about_ca_topic_score_gemma":0.0033173868,"teacher_disagreement_score":0.0022560991,"about_ca_system_score_codex":0.00039573017,"about_ca_system_score_gemma":0.0018364289,"threshold_uncertainty_score":0.011931539},"labels":[],"label_agreement":null},{"id":"W2472565241","doi":"10.1016/j.jss.2016.06.069","title":"Introduction to the special issue on technical debt in software systems","year":2016,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Scalability; Heap (data structure); Abstraction; Transitive relation; Context (archaeology); Software; Computation; Distributed computing; Algorithm; Programming language; Mathematics","score_opus":0.011428635527294556,"score_gpt":0.24877326356731988,"score_spread":0.23734462804002532,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2472565241","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005257887,0.06438581,0.004076374,0.10267697,0.798562,0.00004774847,0.00061185396,0.00020870159,0.028904706],"genre_scores_gemma":[0.002777458,0.027590342,0.0009621669,0.018965254,0.9057474,0.00005482607,0.00059723586,0.00033557904,0.042969767],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982407,0.0003199342,0.00022205939,0.00032421495,0.0007208261,0.00017234682],"domain_scores_gemma":[0.98801583,0.005977002,0.00074973225,0.00076206116,0.0029085614,0.0015868471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00304317,0.0020907223,0.0024206787,0.005457737,0.0018729835,0.0076292166,0.0018358328,0.0051651257,0.053471275],"category_scores_gemma":[0.012042569,0.0007696431,0.0016495412,0.003781182,0.0017911069,0.008220867,0.0030980175,0.009843715,0.01823523],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011099925,0.000019166497,0.00010696211,0.0001694065,0.000009436694,0.000041707586,0.000026662401,0.000078274854,0.00008305681,0.004238884,0.97851425,0.016701039],"study_design_scores_gemma":[0.000009212931,0.000031878124,0.0008022292,0.00044970863,0.00001724266,0.00016596036,0.000050099272,0.00023193279,0.00007260421,0.009302707,0.988848,0.00001853537],"about_ca_topic_score_codex":0.0011993045,"about_ca_topic_score_gemma":0.0026416164,"teacher_disagreement_score":0.053471275,"about_ca_system_score_codex":0.0019015498,"about_ca_system_score_gemma":0.0021697783,"threshold_uncertainty_score":0.1788792},"labels":[],"label_agreement":null},{"id":"W2472584751","doi":"10.1109/tse.2016.2586066","title":"A Study of Causes and Consequences of Client-Side JavaScript Bugs","year":2016,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Intel Corporation","keywords":"Unobtrusive JavaScript; JavaScript; Computer science; Rich Internet application; Document Object Model; Web application; Programmer; World Wide Web; Client-side; Ajax; Programming language; Web page","score_opus":0.024017693639851897,"score_gpt":0.2549549340538039,"score_spread":0.23093724041395203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2472584751","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9948184,0.0007191637,0.0025437356,0.00024047833,0.000018342153,0.000119740995,0.00023879086,0.0002157596,0.0010855797],"genre_scores_gemma":[0.99671066,0.0003347187,0.0020822003,0.00007110564,0.000017470165,0.00005138793,0.00034122722,0.000046845456,0.00034443405],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9901273,0.00210985,0.0010746176,0.001285898,0.00475317,0.0006492412],"domain_scores_gemma":[0.7997758,0.120363146,0.049148068,0.00515491,0.023321213,0.0022369702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046360707,0.0007427322,0.000613566,0.0077504274,0.0010631179,0.0011773563,0.0010202689,0.0010723185,0.0011777787],"category_scores_gemma":[0.0711334,0.00062970276,0.00080470985,0.0043458925,0.0011537491,0.002022413,0.0011139181,0.0010103182,0.0002711611],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023428966,0.00042760622,0.9391607,0.0004607798,0.000107321895,0.0025899021,0.0036962444,0.0014042608,0.003003351,0.00046808913,0.0007878635,0.047659647],"study_design_scores_gemma":[0.00002830602,0.00050414307,0.9792667,0.00018834067,0.00019822763,0.0036130643,0.0037411586,0.0073125735,0.0033617888,0.00054046384,0.0011988371,0.000046508434],"about_ca_topic_score_codex":0.0056458293,"about_ca_topic_score_gemma":0.0061083594,"teacher_disagreement_score":0.0077504274,"about_ca_system_score_codex":0.0012126681,"about_ca_system_score_gemma":0.0021552069,"threshold_uncertainty_score":0.024518132},"labels":[],"label_agreement":null},{"id":"W2474835145","doi":"10.1109/tse.2016.2584050","title":"An Empirical Comparison of Model Validation Techniques for Defect Prediction Models","year":2016,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":566,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; Japan Society for the Promotion of Science; Japan Society for the Promotion of Science London; Compute Canada","keywords":"Computer science; Variance (accounting); Context (archaeology); Cross-validation; Model validation; Sample (material); Data mining; Predictive modelling; Software bug; Software; Machine learning","score_opus":0.054222375516789025,"score_gpt":0.33031650944663127,"score_spread":0.2760941339298422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2474835145","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6999274,0.011148597,0.2750353,0.0014983312,0.00035505716,0.0005150843,0.0033843287,0.002021749,0.0061141704],"genre_scores_gemma":[0.9277461,0.001358534,0.06456732,0.00020037802,0.000074612486,0.0003637347,0.004694594,0.00037236148,0.000622304],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.953718,0.030661775,0.002978476,0.0045603006,0.007287162,0.0007942178],"domain_scores_gemma":[0.51112753,0.4312624,0.012200539,0.025861066,0.018464979,0.001083364],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.088046685,0.0020123732,0.0012534371,0.0054560574,0.001089895,0.0024070835,0.0021225763,0.0025900737,0.001157252],"category_scores_gemma":[0.26811898,0.00067469577,0.0030256722,0.004156694,0.0016456712,0.004996179,0.0021404293,0.003260061,0.0005979206],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022089372,0.000835041,0.26697707,0.002063969,0.0044846353,0.0002960939,0.0015419234,0.45862967,0.0025318018,0.015495986,0.011001714,0.2339331],"study_design_scores_gemma":[0.00020106128,0.0017242713,0.083267085,0.00085884816,0.0006721755,0.0006702605,0.00064923894,0.88682425,0.0038246422,0.015473312,0.0056385295,0.00019630369],"about_ca_topic_score_codex":0.004775952,"about_ca_topic_score_gemma":0.0060357214,"teacher_disagreement_score":0.91195333,"about_ca_system_score_codex":0.002420923,"about_ca_system_score_gemma":0.0019278629,"threshold_uncertainty_score":0.46564096},"labels":[],"label_agreement":null},{"id":"W2475137645","doi":"10.1145/2932631","title":"Multi-Criteria Code Refactoring Using Search-Based Software Engineering","year":2016,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":137,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Japan Society for the Promotion of Science","keywords":"Code refactoring; Computer science; Consistency (knowledge bases); Software quality; Benchmark (surveying); Software engineering; Search-based software engineering; Code smell; Software; Class (philosophy); Source code; Code (set theory); Software development; Software design; Programming language; Artificial intelligence","score_opus":0.17718848535030324,"score_gpt":0.36867093108150445,"score_spread":0.1914824457312012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2475137645","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09506798,0.0011389497,0.89731044,0.00032815168,0.00005044449,0.00063978916,0.00013948468,0.0024636977,0.0028609827],"genre_scores_gemma":[0.3452206,0.00023940137,0.65166944,0.0001928225,0.000029421792,0.00040380534,0.0004458659,0.00024745724,0.0015511501],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99552625,0.0015426842,0.0003428651,0.00075693603,0.0015991774,0.00023208608],"domain_scores_gemma":[0.9913339,0.0053983815,0.0009318027,0.00057962426,0.0015298134,0.00022644295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046374486,0.002499873,0.0021837344,0.005374388,0.0008700876,0.001124849,0.0026742623,0.0017328786,0.0017606147],"category_scores_gemma":[0.013034788,0.00088885834,0.0017300138,0.0028081175,0.0009062238,0.0012987575,0.0016350774,0.0010096933,0.00048453722],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025128791,0.00069857173,0.0045741033,0.0005219888,0.00028253975,0.0002414539,0.00040974142,0.6370179,0.009899888,0.0037841168,0.0018699066,0.34044853],"study_design_scores_gemma":[0.000060697173,0.00019132493,0.00063988566,0.00003318688,0.00005430854,0.000059035756,0.00005339431,0.99476516,0.0018127164,0.0017244141,0.000587496,0.000018389479],"about_ca_topic_score_codex":0.010444199,"about_ca_topic_score_gemma":0.012172607,"teacher_disagreement_score":0.010444199,"about_ca_system_score_codex":0.0016255851,"about_ca_system_score_gemma":0.002543405,"threshold_uncertainty_score":0.024525464},"labels":[],"label_agreement":null},{"id":"W2480141206","doi":"10.1007/978-3-319-25964-2_6","title":"Measuring the Utility of Functional-Based Software Using Centroid-Adjusted Class Labelling","year":2016,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"RTDS Technologies (Canada); University of Manitoba","funders":"","keywords":"Computer science; Software; Artificial intelligence; Python (programming language); Preprocessor; Class (philosophy); Data mining; Machine learning; Centroid; Programming language","score_opus":0.24767688385340805,"score_gpt":0.34537383138558003,"score_spread":0.09769694753217198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2480141206","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5138879,0.001169224,0.46970567,0.00023148052,0.00009621494,0.000106675376,0.0008601402,0.003466676,0.010475992],"genre_scores_gemma":[0.8558606,0.00015735121,0.14127693,0.00002321101,0.000020281741,0.00004248497,0.00094737374,0.00035684157,0.001314937],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9960375,0.00072023115,0.00017515809,0.0005388405,0.002330423,0.00019785609],"domain_scores_gemma":[0.9886275,0.0059625925,0.00097076193,0.0015013503,0.0026678648,0.0002699641],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023822328,0.00061992934,0.00067787647,0.0041293222,0.00049991824,0.0018952605,0.0017825596,0.0012251934,0.0018899487],"category_scores_gemma":[0.025016865,0.00022021402,0.0006544873,0.0042947466,0.0007721695,0.0030741813,0.0012602727,0.00075602817,0.0009054185],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009212798,0.00026953232,0.034964003,0.00039992074,0.000199896,0.00011613011,0.00074377476,0.09140773,0.02859782,0.011375684,0.0047988608,0.8262054],"study_design_scores_gemma":[0.00003030315,0.00033903273,0.03419818,0.000050033068,0.00009465622,0.0002917859,0.00036474646,0.9238579,0.022278313,0.014869413,0.0035459038,0.00007970853],"about_ca_topic_score_codex":0.0049357014,"about_ca_topic_score_gemma":0.0058249556,"teacher_disagreement_score":0.0049357014,"about_ca_system_score_codex":0.0011738518,"about_ca_system_score_gemma":0.00062878075,"threshold_uncertainty_score":0.012598634},"labels":[],"label_agreement":null},{"id":"W2480481610","doi":"10.1007/978-3-540-71389-0","title":"Foundations of Software Science and Computational Structures","year":2007,"lang":"en","type":"book","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure","funders":"","keywords":"Computer science; Software engineering; Software; Programming language; Computational science; Theoretical computer science","score_opus":0.018231483566501717,"score_gpt":0.29185699549625393,"score_spread":0.2736255119297522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2480481610","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003789264,0.03941832,0.25173703,0.0077035567,0.0024307694,0.00009295865,0.00045643063,0.00094630383,0.69342536],"genre_scores_gemma":[0.15045474,0.07923975,0.19120976,0.0025611338,0.0058923857,0.00056440773,0.0013957708,0.0009571097,0.56772494],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99970895,0.00004121707,0.000014857388,0.000040975567,0.00017053657,0.000023355704],"domain_scores_gemma":[0.99938667,0.00032955784,0.000030354453,0.00010373597,0.00011701295,0.000032656793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000324102,0.0008799352,0.00084055134,0.0014604839,0.0009704895,0.0029457032,0.0008130772,0.00094189774,0.017222678],"category_scores_gemma":[0.0013797968,0.0006012073,0.00062501506,0.0027749755,0.002208415,0.00424387,0.0009134367,0.0026439957,0.0075349575],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000004645636,0.000012653343,0.00004048228,0.00013936695,0.000004941576,0.000019402261,0.00010546439,0.00088037446,0.0003045853,0.90703136,0.03479213,0.05666458],"study_design_scores_gemma":[0.00000464979,0.0000047924973,0.00008341203,0.00006281541,0.000005321187,0.0000718476,0.00002233036,0.0014240418,0.00017064952,0.81493956,0.18320565,0.000004951928],"about_ca_topic_score_codex":0.0009439195,"about_ca_topic_score_gemma":0.0012489283,"teacher_disagreement_score":0.017222678,"about_ca_system_score_codex":0.0013776205,"about_ca_system_score_gemma":0.0016636994,"threshold_uncertainty_score":0.05761558},"labels":[],"label_agreement":null},{"id":"W2481451249","doi":"10.1145/2851613.2851792","title":"Is code cloning in games really different?","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Cloning (programming); Computer science; Source code; Programming language; Reuse; Java; Code (set theory); Software; Code reuse; clone (Java method); Game programming; Visualization; Software visualization; Theoretical computer science; Video game development; Game design; Game Developer; Human–computer interaction; Software system; Artificial intelligence; Component-based software engineering; Engineering; Set (abstract data type); Biology","score_opus":0.022576336881930326,"score_gpt":0.2790514103561447,"score_spread":0.2564750734742144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2481451249","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9333726,0.0024761974,0.048508942,0.005293402,0.00019924082,0.00020424754,0.00023336259,0.00052331114,0.009188769],"genre_scores_gemma":[0.97822535,0.0008604356,0.016789597,0.0015708382,0.000058666003,0.00008861285,0.00027555437,0.00029708364,0.0018339992],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9678684,0.0123678595,0.0023179224,0.0057369457,0.010165518,0.0015432904],"domain_scores_gemma":[0.82107824,0.1085752,0.030288491,0.021591606,0.01530901,0.0031574443],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0122560095,0.0006536009,0.0008250342,0.0028019634,0.0018418885,0.0053587346,0.0021556325,0.0019857904,0.0016128548],"category_scores_gemma":[0.17915003,0.00071647036,0.0011115596,0.0031118826,0.009056356,0.012283973,0.003154606,0.002317416,0.0005485472],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061274215,0.0004933929,0.582976,0.0014099469,0.0006551052,0.0015516111,0.048144586,0.0031347862,0.011965209,0.047295246,0.0049194735,0.29684195],"study_design_scores_gemma":[0.00012715478,0.0015071038,0.721573,0.0020556492,0.0011100182,0.00899234,0.056373905,0.019060455,0.027190685,0.10105199,0.06052157,0.00043607684],"about_ca_topic_score_codex":0.0074816747,"about_ca_topic_score_gemma":0.008057292,"teacher_disagreement_score":0.0122560095,"about_ca_system_score_codex":0.0026983146,"about_ca_system_score_gemma":0.0026318305,"threshold_uncertainty_score":0.06481677},"labels":[],"label_agreement":null},{"id":"W2484350885","doi":"10.4018/978-1-59140-896-3.ch005","title":"Design Patterns as Laws of Quality","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Quality (philosophy); Computer science; Software quality; Software; Software quality control; Object (grammar); Data mining; Measure (data warehouse); Software engineering; Artificial intelligence; Software development; Programming language","score_opus":0.06916189357925827,"score_gpt":0.302254467090375,"score_spread":0.23309257351111673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2484350885","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02248622,0.010232403,0.62784094,0.016666906,0.00047525746,0.00027272216,0.00024350289,0.00077041733,0.3210117],"genre_scores_gemma":[0.44731298,0.011704144,0.45187688,0.00312908,0.00051899237,0.0015112028,0.00045497032,0.00071014947,0.082781695],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99562573,0.001762654,0.0003796095,0.0006688308,0.0013330886,0.00023011661],"domain_scores_gemma":[0.9919629,0.004393705,0.0006582131,0.0019401317,0.00082527724,0.00021977744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004247091,0.00085917354,0.00057804905,0.0019503135,0.0015968127,0.006784692,0.0015825558,0.0029553038,0.0061829085],"category_scores_gemma":[0.011829525,0.001064812,0.00090560445,0.0022181505,0.01341516,0.010982503,0.002655818,0.003772726,0.0014358979],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000024016008,0.000006457694,0.00012701192,0.000054766006,0.0000047015187,0.000025574458,0.0005074782,0.0011851855,0.00012231414,0.9862472,0.0014938264,0.010223064],"study_design_scores_gemma":[0.00000756185,0.0000071532327,0.00008518895,0.00006572642,0.0000047006088,0.000059076512,0.000079694706,0.0021071192,0.00014853844,0.9665896,0.030840514,0.0000051068787],"about_ca_topic_score_codex":0.0020143392,"about_ca_topic_score_gemma":0.0015204838,"teacher_disagreement_score":0.006784692,"about_ca_system_score_codex":0.0039596893,"about_ca_system_score_gemma":0.0018808865,"threshold_uncertainty_score":0.028729677},"labels":[],"label_agreement":null},{"id":"W2484574729","doi":"10.1016/b978-0-12-804206-9.00022-2","title":"What counts is decisions, not numbers—Toward an analytics design sheet","year":2016,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Analytics; Computer science; Data science","score_opus":0.06571567316967893,"score_gpt":0.29790900435180284,"score_spread":0.23219333118212393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2484574729","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002608839,0.009555308,0.696702,0.14228585,0.004909536,0.00023613081,0.0007401455,0.0012662234,0.14169592],"genre_scores_gemma":[0.11572205,0.021070572,0.7298948,0.025925145,0.007260728,0.0013138329,0.0012300609,0.0015663584,0.096016444],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99484396,0.0025077544,0.00061753375,0.0005517496,0.0013423194,0.0001366738],"domain_scores_gemma":[0.9809564,0.0136989225,0.00075264904,0.0013252245,0.0026207187,0.0006461968],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009155735,0.0018752337,0.0013547377,0.0041408427,0.002611455,0.015045322,0.0023820174,0.002930731,0.009882761],"category_scores_gemma":[0.01955311,0.0012585854,0.0008722321,0.003468012,0.014794586,0.030479684,0.0029098934,0.008551492,0.006241283],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013142037,0.0000145705935,0.0002172836,0.00015602246,0.0000079929105,0.000024046687,0.0008818784,0.00048503967,0.00019377629,0.9222348,0.033503897,0.042267438],"study_design_scores_gemma":[0.0000049368546,0.000012625877,0.000092245995,0.0002652691,0.0000069584694,0.000047246485,0.00054875267,0.0015573321,0.00020714082,0.85498023,0.14226198,0.0000151757085],"about_ca_topic_score_codex":0.002309562,"about_ca_topic_score_gemma":0.002330496,"teacher_disagreement_score":0.015045322,"about_ca_system_score_codex":0.0030829168,"about_ca_system_score_gemma":0.0048227934,"threshold_uncertainty_score":0.048420727},"labels":[],"label_agreement":null},{"id":"W2485677833","doi":"10.1007/978-3-319-41579-6_13","title":"Estimating Development Effort for Software Architectural Tactics","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Computer Research Institute of Montréal","funders":"","keywords":"Computer science; Analytic hierarchy process; Key (lock); Quality (philosophy); Software; Software engineering; Process (computing); Software development; Product (mathematics); Software quality; Estimation; Hierarchy; Risk analysis (engineering); Systems engineering; Operations research; Computer security; Engineering","score_opus":0.020850553755454176,"score_gpt":0.2705426495716893,"score_spread":0.2496920958162351,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2485677833","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55707633,0.0022200642,0.4291914,0.00016882094,0.000051071358,0.00013926871,0.0009881769,0.0024420242,0.007722883],"genre_scores_gemma":[0.7084715,0.0006215853,0.28559372,0.000016599373,0.000028473041,0.0001217806,0.002062482,0.00031042416,0.002773541],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976755,0.00064491434,0.00017334394,0.00035886472,0.0009883092,0.00015905993],"domain_scores_gemma":[0.97816205,0.016410165,0.0014478012,0.0016570258,0.002066549,0.00025649247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018425647,0.0011103341,0.00049635425,0.004365117,0.00030709623,0.0008222437,0.00088851666,0.0006308961,0.0020572061],"category_scores_gemma":[0.024757326,0.0005254834,0.00091498985,0.0024713357,0.00018932908,0.0019633006,0.0007005489,0.0006880019,0.0006350366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003523468,0.0002432957,0.081656136,0.00034563974,0.0001487631,0.00011207638,0.00028804384,0.15389104,0.008260701,0.0035582203,0.0019834323,0.7491603],"study_design_scores_gemma":[0.00003270782,0.00041893168,0.040984094,0.00008525435,0.00014812124,0.00018995145,0.00017077503,0.93893105,0.010418348,0.0066353115,0.001952845,0.000032624266],"about_ca_topic_score_codex":0.0031553903,"about_ca_topic_score_gemma":0.005355772,"teacher_disagreement_score":0.004365117,"about_ca_system_score_codex":0.0007704502,"about_ca_system_score_gemma":0.00065968314,"threshold_uncertainty_score":0.009744525},"labels":[],"label_agreement":null},{"id":"W2486735965","doi":"10.4018/978-1-60566-060-8.ch018","title":"Intelligent Analysis of Software Maintenance Data","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software; Software maintenance; Software development; Software engineering; Process (computing); Data mining; Software development process; Data extraction; Software construction; Software sizing","score_opus":0.03945904192212604,"score_gpt":0.28856905740952044,"score_spread":0.2491100154873944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2486735965","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27893063,0.0014954903,0.6945673,0.00083208166,0.00008920784,0.00047855213,0.0070754196,0.010557014,0.0059742047],"genre_scores_gemma":[0.53568226,0.00074731535,0.4473938,0.00012218147,0.00006107911,0.0002811403,0.013561464,0.00022697976,0.0019237661],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977228,0.00041266045,0.00024840457,0.00053992006,0.0009760297,0.00010009744],"domain_scores_gemma":[0.9945486,0.002622341,0.0005744847,0.00078137,0.0013891419,0.000083974905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024281566,0.00056775165,0.0009697506,0.0058657555,0.00043606682,0.0017239371,0.0007617268,0.0005264515,0.00092211424],"category_scores_gemma":[0.00991506,0.00030697678,0.00088584085,0.002870682,0.00021631314,0.001431081,0.0009661926,0.00079829356,0.0005669492],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003132749,0.00019347263,0.03649259,0.000413201,0.00020107249,0.0003777242,0.0006819973,0.057738956,0.017734235,0.004106125,0.0050648367,0.8766826],"study_design_scores_gemma":[0.000017940492,0.0001481746,0.03940437,0.000107438456,0.000110239045,0.00029643727,0.00033676572,0.91611475,0.020437518,0.009901074,0.0130737005,0.000051585386],"about_ca_topic_score_codex":0.0022624603,"about_ca_topic_score_gemma":0.0031506545,"teacher_disagreement_score":0.0058657555,"about_ca_system_score_codex":0.0007019225,"about_ca_system_score_gemma":0.00080609554,"threshold_uncertainty_score":0.012841463},"labels":[],"label_agreement":null},{"id":"W2489048129","doi":"10.4018/978-1-4666-4301-7.ch024","title":"A Framework for Testing Code in Computational Applications","year":2013,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Apheresis Group; Royal Military College of Canada","funders":"","keywords":"Computer science; Software engineering; Mindset; Software testing; Code (set theory); Process (computing); Computational model; Test strategy; Software; Programming language; Set (abstract data type); Simulation; Artificial intelligence","score_opus":0.04577002947142524,"score_gpt":0.3058061314727236,"score_spread":0.2600361020012984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2489048129","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012941383,0.009344149,0.8229592,0.009801222,0.000729886,0.00031207612,0.0003125613,0.0012205207,0.15402627],"genre_scores_gemma":[0.069885455,0.011720044,0.8503805,0.0037730613,0.00097489945,0.001523483,0.00067478477,0.00094290066,0.060124952],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9956436,0.002033222,0.00036997793,0.0005090912,0.0012197454,0.00022432762],"domain_scores_gemma":[0.9937018,0.0042501106,0.00024008201,0.0008152367,0.0007805111,0.00021223645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006028171,0.0021438764,0.0008822095,0.0044494313,0.00201372,0.00589382,0.0044034687,0.004726924,0.014844084],"category_scores_gemma":[0.009272476,0.0010535351,0.0018978385,0.0033394028,0.010093583,0.010743952,0.003799389,0.0073431632,0.004751286],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000021002902,0.000009677842,0.000032938115,0.00009220374,0.0000026414534,0.00006757669,0.0002452137,0.001656756,0.00014555601,0.97837514,0.004615495,0.014754717],"study_design_scores_gemma":[0.000005345364,0.000018521587,0.00007039632,0.00041089667,0.0000051507595,0.00021337486,0.00014304166,0.0064917603,0.00030225024,0.8305457,0.16177651,0.000017078693],"about_ca_topic_score_codex":0.004183636,"about_ca_topic_score_gemma":0.0034570554,"teacher_disagreement_score":0.014844084,"about_ca_system_score_codex":0.0042827246,"about_ca_system_score_gemma":0.0033279574,"threshold_uncertainty_score":0.049658418},"labels":[],"label_agreement":null},{"id":"W2489675867","doi":"10.4018/978-1-4666-8111-8.ch022","title":"Software Security Engineering – Part I","year":2015,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; University of Victoria","funders":"","keywords":"Software security assurance; Software development; Social software engineering; Computer science; Personal software process; Software construction; Software peer review; Security bug; Software development process; Security engineering; Software engineering; Package development process; Software; Computer security; Security service; Information security; Operating system","score_opus":0.025682568175708322,"score_gpt":0.2525117060764229,"score_spread":0.2268291379007146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2489675867","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00460582,0.43025842,0.08883681,0.006565678,0.0054779323,0.00026331627,0.00039680325,0.0008257461,0.4627694],"genre_scores_gemma":[0.05115863,0.45073178,0.05407254,0.0045217923,0.0048346263,0.00039235872,0.0014727112,0.00082020415,0.43199545],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99922884,0.00013028645,0.00006861084,0.00015986963,0.00036664528,0.00004570742],"domain_scores_gemma":[0.99942005,0.0002531663,0.000041359253,0.0000973676,0.00014491122,0.000043262025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006172645,0.0012518592,0.0006410827,0.0023018443,0.0007271629,0.0030736553,0.0007254937,0.0014645514,0.017604308],"category_scores_gemma":[0.0012985565,0.00060463935,0.0006836564,0.0030025025,0.0017284721,0.003323183,0.00149953,0.0026425638,0.012138579],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020812888,0.00014121796,0.00036750388,0.002026448,0.000031174022,0.0001722875,0.0012327314,0.0031855402,0.0032823312,0.19385989,0.13957024,0.65610987],"study_design_scores_gemma":[0.0000034534705,0.000047332203,0.0005763823,0.0012186175,0.000006894782,0.00048137933,0.00014694784,0.0007524358,0.00072701223,0.077124685,0.9188996,0.000015284757],"about_ca_topic_score_codex":0.0008121259,"about_ca_topic_score_gemma":0.0008473861,"teacher_disagreement_score":0.017604308,"about_ca_system_score_codex":0.001374955,"about_ca_system_score_gemma":0.0014056192,"threshold_uncertainty_score":0.05889231},"labels":[],"label_agreement":null},{"id":"W2494647806","doi":"10.4018/978-1-60566-026-4.ch098","title":"Communicability of Natural Language in Software Representations","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Natural language; Syntax; Formality; Programming language; Software; Software development; Linguistics; Software engineering; Natural language processing","score_opus":0.01835975397319349,"score_gpt":0.29522967537732936,"score_spread":0.27686992140413585,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2494647806","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016538989,0.018366197,0.70752084,0.013406374,0.00068232155,0.000120312114,0.00014761166,0.00057761,0.24263977],"genre_scores_gemma":[0.5359076,0.019795435,0.38443974,0.0029547696,0.0011885756,0.00071040133,0.0006510361,0.00066045206,0.053691883],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99556255,0.002654496,0.00026390585,0.00047825745,0.00089793507,0.000142847],"domain_scores_gemma":[0.99207604,0.006244578,0.00028504926,0.0009826821,0.00031523683,0.00009636873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004446207,0.0008918164,0.0006339937,0.0021267636,0.0014220387,0.008154671,0.0017475355,0.002974024,0.0047070924],"category_scores_gemma":[0.008697053,0.0005458047,0.00093858445,0.00240666,0.013478183,0.017742813,0.004083657,0.0043464084,0.0010853048],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000033293288,0.0000046498603,0.00003830034,0.000059204296,0.0000025941401,0.0000366226,0.0019205688,0.00035545707,0.00019375166,0.98553896,0.0007628388,0.011083596],"study_design_scores_gemma":[0.0000059319714,0.000012887273,0.00008988212,0.00018374182,0.000008158149,0.00021196772,0.0006362981,0.0024913717,0.00048654614,0.91892093,0.076936245,0.000015905434],"about_ca_topic_score_codex":0.0012604707,"about_ca_topic_score_gemma":0.00090822484,"teacher_disagreement_score":0.008154671,"about_ca_system_score_codex":0.003139151,"about_ca_system_score_gemma":0.0016828609,"threshold_uncertainty_score":0.023514092},"labels":[],"label_agreement":null},{"id":"W2495532368","doi":"10.4018/978-1-59140-941-1.ch008","title":"Modeling Relevance Relations Using Machine Learning Techniques","year":2007,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Relevance (law); Computer science; Artificial intelligence; Relation (database); Machine learning; Software; Software deployment; Abstraction; Precision and recall; Data mining; Data science; Software engineering; Programming language","score_opus":0.04374587496380003,"score_gpt":0.29995860257629564,"score_spread":0.2562127276124956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2495532368","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018080955,0.0028821386,0.9706194,0.0011718465,0.000087611064,0.00018303824,0.00048032252,0.001013813,0.0054808483],"genre_scores_gemma":[0.3614319,0.0030695535,0.62575316,0.00046663635,0.00043570995,0.0005729354,0.002051138,0.00033422213,0.0058847307],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9939017,0.002333607,0.00041306898,0.0011691041,0.0018906653,0.00029191337],"domain_scores_gemma":[0.97978485,0.016794663,0.0011100504,0.0010489239,0.0011012102,0.00016039844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056442283,0.0014874529,0.0015927926,0.0064152605,0.001308961,0.003680916,0.0027305954,0.0019498466,0.0037427293],"category_scores_gemma":[0.0291829,0.00096940505,0.0019816717,0.00519699,0.0013074802,0.008053298,0.0019629803,0.0031603316,0.0015617211],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016196776,0.00037230976,0.010785502,0.00065564556,0.00027074257,0.00086503057,0.0011230184,0.4063152,0.002264223,0.15883854,0.015967075,0.4023808],"study_design_scores_gemma":[0.000013280005,0.000025487678,0.0009460145,0.00007548531,0.000037834605,0.0001942629,0.00007297277,0.839862,0.0006457201,0.1516599,0.0064426446,0.000024390145],"about_ca_topic_score_codex":0.006202288,"about_ca_topic_score_gemma":0.00653259,"teacher_disagreement_score":0.0064152605,"about_ca_system_score_codex":0.0022232875,"about_ca_system_score_gemma":0.0014807972,"threshold_uncertainty_score":0.029849887},"labels":[],"label_agreement":null},{"id":"W2495954116","doi":"10.1016/b978-0-12-804206-9.00017-9","title":"A success story in applying data science in practice","year":2016,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Data science","score_opus":0.0383235574101892,"score_gpt":0.3130281940902187,"score_spread":0.2747046366800295,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2495954116","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0073481547,0.06855994,0.05160126,0.5389331,0.007562018,0.00009570191,0.0001773037,0.00041787585,0.32530466],"genre_scores_gemma":[0.32080245,0.10793341,0.09438631,0.064005814,0.012296658,0.00072004297,0.00035460165,0.002353948,0.39714682],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9839767,0.007970021,0.00069511775,0.00089829374,0.0058911373,0.000568734],"domain_scores_gemma":[0.9463077,0.043280333,0.0008417717,0.0037604664,0.0035454768,0.0022642824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026593983,0.0009539234,0.0010379842,0.0033587264,0.0050757825,0.026957652,0.001839826,0.0069018146,0.011931634],"category_scores_gemma":[0.041980665,0.00089596346,0.00068333664,0.0062786667,0.034255482,0.036020003,0.011827462,0.014550838,0.004215608],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020082547,0.000027023916,0.00025701147,0.0002636783,0.000009404784,0.00013757503,0.005730525,0.0001517683,0.00013863077,0.8547673,0.07206069,0.06643633],"study_design_scores_gemma":[0.00001754481,0.000025660082,0.00016474539,0.00064185454,0.0000060751713,0.00018790235,0.0040279143,0.00047585994,0.00026913267,0.3078843,0.68628055,0.000018543386],"about_ca_topic_score_codex":0.0035271926,"about_ca_topic_score_gemma":0.004166294,"teacher_disagreement_score":0.026957652,"about_ca_system_score_codex":0.006490163,"about_ca_system_score_gemma":0.006467481,"threshold_uncertainty_score":0.14064413},"labels":[],"label_agreement":null},{"id":"W2498728262","doi":"10.4018/978-1-61520-965-1.ch308","title":"Integrating Software Engineering and Costing Aspects within Project Management Tools","year":2010,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Gantt chart; Software engineering; Computer science; Activity-based costing; Systems engineering; Software project management; Software; Software development; Engineering management; Engineering drawing; Engineering; Software construction; Programming language","score_opus":0.021049591290809827,"score_gpt":0.25067058998093317,"score_spread":0.22962099869012334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2498728262","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027546246,0.00067891617,0.94274247,0.0005184539,0.000059351394,0.0001263308,0.00012621183,0.001598561,0.02660342],"genre_scores_gemma":[0.2839331,0.0016248983,0.70218784,0.00012880877,0.000042795935,0.00022310478,0.0004987186,0.00049167965,0.010869112],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99757475,0.0008626391,0.00016397139,0.0003234766,0.00095661037,0.00011852163],"domain_scores_gemma":[0.9965661,0.001971564,0.00025297646,0.00077808916,0.00032990985,0.00010138024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002063468,0.00094561046,0.00054034963,0.002438617,0.0005542854,0.0054937745,0.0012918297,0.0010171492,0.0028498082],"category_scores_gemma":[0.0052370704,0.0006170301,0.00078083575,0.0029380708,0.0010306686,0.007277483,0.001455589,0.0011940553,0.0010691483],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006767139,0.00020874712,0.0056550624,0.00056006905,0.00005172997,0.00033718732,0.0023189127,0.058891825,0.0062366514,0.30285972,0.0030346918,0.6197778],"study_design_scores_gemma":[0.000043157834,0.00028959423,0.0078056073,0.0011779795,0.00015227312,0.0014438159,0.0013982664,0.32773504,0.014368843,0.32678467,0.31863743,0.0001633716],"about_ca_topic_score_codex":0.0030163801,"about_ca_topic_score_gemma":0.0026947586,"teacher_disagreement_score":0.0054937745,"about_ca_system_score_codex":0.0012125053,"about_ca_system_score_gemma":0.0021691725,"threshold_uncertainty_score":0.010912776},"labels":[],"label_agreement":null},{"id":"W2500014933","doi":"10.4018/978-1-60566-006-6.ch013","title":"A Framework for Understanding and Addressing the Semiotic Quality of Use Case Models","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Semiotics; Scope (computer science); Computer science; Quality (philosophy); Process (computing); Product (mathematics); Management science; Process management; Risk analysis (engineering); Data science; Knowledge management; Engineering; Epistemology; Business","score_opus":0.2820967381977312,"score_gpt":0.3723759483449482,"score_spread":0.09027921014721701,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2500014933","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012083027,0.0035697592,0.9409649,0.005695577,0.00019634425,0.0003661366,0.00011771833,0.00022472926,0.047656607],"genre_scores_gemma":[0.06018821,0.004575922,0.9231426,0.00077076914,0.0002591136,0.0014422499,0.0003883809,0.0001604084,0.009072297],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97853756,0.014905269,0.0017486546,0.0009055806,0.0034482856,0.00045462057],"domain_scores_gemma":[0.96882576,0.023524744,0.0015417424,0.0027831472,0.002761266,0.0005632974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025123026,0.0018342207,0.0013276101,0.009232109,0.0033797398,0.014500587,0.0045129107,0.005163656,0.006244785],"category_scores_gemma":[0.026602626,0.0014095347,0.002051617,0.007064571,0.021164734,0.017341927,0.0064373063,0.007395827,0.001637471],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000021494498,0.000008024745,0.00004540253,0.00009280224,0.0000042560205,0.000071229304,0.0015998015,0.000737256,0.00010589174,0.9884271,0.0009742793,0.00793178],"study_design_scores_gemma":[0.000006477339,0.0000124772805,0.00006997488,0.000462789,0.0000097362945,0.00021917792,0.0010801101,0.0043984707,0.00023536076,0.9152674,0.078218915,0.000019211553],"about_ca_topic_score_codex":0.0040345006,"about_ca_topic_score_gemma":0.003498102,"teacher_disagreement_score":0.025123026,"about_ca_system_score_codex":0.008434697,"about_ca_system_score_gemma":0.0071704932,"threshold_uncertainty_score":0.13286489},"labels":[],"label_agreement":null},{"id":"W2504374816","doi":"10.4018/978-1-59904-090-5.ch016","title":"Software Specification and Attack Languages","year":2007,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Programming language; Software engineering; Software security assurance; Software development; Software requirements specification; Specification language; Software construction; Software; Computer security; Information security; Security service","score_opus":0.03964932326458692,"score_gpt":0.3014450790820726,"score_spread":0.26179575581748565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2504374816","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051365546,0.0068029184,0.93086386,0.0030390765,0.0005458523,0.00047275907,0.0009134996,0.0032173395,0.04900821],"genre_scores_gemma":[0.13973103,0.017823922,0.7824979,0.004169089,0.0009935695,0.001777907,0.005827685,0.0020840284,0.045094907],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9857335,0.0051201005,0.0025712587,0.0014491035,0.004327251,0.00079883815],"domain_scores_gemma":[0.98849267,0.005499822,0.0012824908,0.0020727064,0.002398261,0.0002540479],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064720693,0.0020990258,0.0010150156,0.003169594,0.001584237,0.006144291,0.00213079,0.0028961047,0.008086043],"category_scores_gemma":[0.012267982,0.0011248223,0.001949726,0.0041622855,0.004893886,0.008748042,0.0034250324,0.0052088173,0.0047965734],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003679223,0.00003415737,0.00031561073,0.0006472031,0.000021717937,0.00021520308,0.0014339418,0.0028177616,0.0019997554,0.9192224,0.01111592,0.06213945],"study_design_scores_gemma":[0.000035541987,0.00009242253,0.00024266509,0.0006712737,0.000037112517,0.0014844178,0.00056465896,0.012346525,0.004453256,0.3043661,0.6756428,0.000063114145],"about_ca_topic_score_codex":0.002436226,"about_ca_topic_score_gemma":0.0012983574,"teacher_disagreement_score":0.008086043,"about_ca_system_score_codex":0.0024742773,"about_ca_system_score_gemma":0.0035795663,"threshold_uncertainty_score":0.034227967},"labels":[],"label_agreement":null},{"id":"W2504514184","doi":"10.1109/wcre.1996.558936","title":"A cliche-based environment to support architectural reverse engineering","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Software engineering; Architectural pattern; Architectural geometry; Software; Software architecture; Architectural model; Architecture; Reverse engineering; Software system; Set (abstract data type); Cliché; Programming language; Systems engineering; Software construction; Engineering","score_opus":0.020317222176899862,"score_gpt":0.21848293671885802,"score_spread":0.19816571454195817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2504514184","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003627447,0.0002118793,0.8002036,0.00043784248,0.00029750084,0.0004475151,0.0021867075,0.1785632,0.014024369],"genre_scores_gemma":[0.039406557,0.00034785861,0.8896808,0.00055279135,0.00015040828,0.001153482,0.009405135,0.024364177,0.034938727],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99841964,0.00033241857,0.00015484632,0.00027482616,0.00070455135,0.000113678674],"domain_scores_gemma":[0.99500954,0.001752299,0.00025695213,0.0018322415,0.0007359609,0.00041297328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028246965,0.0014049608,0.001058506,0.002196994,0.0013872822,0.003671506,0.004064011,0.0022529243,0.04613578],"category_scores_gemma":[0.009350407,0.0014254649,0.001213076,0.0016147017,0.0011183865,0.0057351985,0.007385945,0.0037851376,0.01716492],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018355937,0.0007473846,0.0022260305,0.00088546315,0.00013825568,0.0024255617,0.0020152114,0.021186613,0.029307317,0.09070847,0.28972295,0.55880123],"study_design_scores_gemma":[0.0005782392,0.00022705262,0.00063529745,0.00020718949,0.000059133872,0.00086393114,0.00019887024,0.16040751,0.02034571,0.065571845,0.7507007,0.00020450303],"about_ca_topic_score_codex":0.0016004964,"about_ca_topic_score_gemma":0.0032135046,"teacher_disagreement_score":0.04613578,"about_ca_system_score_codex":0.0006837731,"about_ca_system_score_gemma":0.0011465636,"threshold_uncertainty_score":0.15433955},"labels":[],"label_agreement":null},{"id":"W2507266217","doi":"10.4995/thesis/10251/66868","title":"Detección de reutilización de código fuente monolingüe y translingüe","year":2016,"lang":"es","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Source code; Set (abstract data type); Temptation; Suspect; Programming language; Code (set theory); Software; Source lines of code","score_opus":0.01789252567758166,"score_gpt":0.3039440512791934,"score_spread":0.28605152560161173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2507266217","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36827463,0.0062830425,0.5475759,0.0063990876,0.00077371125,0.0008225093,0.0012954619,0.020314068,0.048261553],"genre_scores_gemma":[0.562713,0.0030423745,0.39215985,0.00077057973,0.00014359213,0.00033947683,0.0021858476,0.007837315,0.030807922],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99241054,0.0030138388,0.0006259158,0.0014110061,0.0018894316,0.0006492512],"domain_scores_gemma":[0.97138095,0.011809471,0.001313096,0.007405114,0.0072425823,0.00084875635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009348706,0.0013750985,0.0010855519,0.002301936,0.001591123,0.0075391233,0.0015471876,0.0014983066,0.0066473507],"category_scores_gemma":[0.036093578,0.0012297531,0.00123755,0.0012832487,0.00222693,0.011832724,0.00538725,0.003112069,0.003198236],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013097128,0.0004681005,0.018275807,0.0015274538,0.0001901517,0.0022140336,0.02887797,0.007089189,0.07726208,0.055268582,0.023754848,0.7837621],"study_design_scores_gemma":[0.00018293879,0.00077533844,0.01436329,0.0021361324,0.0008613487,0.004133378,0.021802425,0.17159992,0.21732107,0.062239185,0.5041817,0.00040335717],"about_ca_topic_score_codex":0.009975014,"about_ca_topic_score_gemma":0.012753848,"teacher_disagreement_score":0.009975014,"about_ca_system_score_codex":0.0021517598,"about_ca_system_score_gemma":0.005196125,"threshold_uncertainty_score":0.049441278},"labels":[],"label_agreement":null},{"id":"W2507516899","doi":"10.1007/s10270-016-0557-6","title":"An approach to clone detection in sequence diagrams and its application to security analysis","year":2016,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Sequence diagram; clone (Java method); Sequence (biology); Software engineering; Programming language; Unified Modeling Language; Software","score_opus":0.027096577173787176,"score_gpt":0.2782843791539248,"score_spread":0.2511878019801376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2507516899","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022816313,0.000040460895,0.99645656,0.000064264554,0.00000872957,0.000045019988,0.00002297362,0.00077109487,0.00030923277],"genre_scores_gemma":[0.057212785,0.00012548658,0.9411016,0.00006126883,0.000022449833,0.0000760031,0.00010667545,0.00018815917,0.0011055534],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969909,0.000773454,0.0002579177,0.000522775,0.0013083997,0.00014646472],"domain_scores_gemma":[0.9890615,0.006381294,0.0007013184,0.0015692585,0.0019756407,0.00031112062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028477951,0.0009752631,0.0008896141,0.0055809706,0.0015834626,0.0023388204,0.0017425225,0.002353664,0.0026796856],"category_scores_gemma":[0.013270029,0.0009648592,0.0019502224,0.0031801807,0.0019993342,0.0039344747,0.0023458079,0.002586387,0.0009039488],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018699239,0.00045598362,0.008202599,0.00048327423,0.00016622648,0.0011471958,0.0019778598,0.061052322,0.040322825,0.29500222,0.0042593544,0.5867431],"study_design_scores_gemma":[0.000040669023,0.00013820986,0.0014341603,0.00014695649,0.00018044627,0.001655126,0.0003216741,0.7544716,0.024055867,0.2025781,0.014886616,0.00009061466],"about_ca_topic_score_codex":0.004668957,"about_ca_topic_score_gemma":0.0055856565,"teacher_disagreement_score":0.0055809706,"about_ca_system_score_codex":0.0010771151,"about_ca_system_score_gemma":0.0019294736,"threshold_uncertainty_score":0.015060723},"labels":[],"label_agreement":null},{"id":"W2510065047","doi":"10.18293/seke2016-250","title":"Clustering and Artificial Neural Network Ensembles Based Effort Estimation","year":2016,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Menofia University; University of Calgary","keywords":"Cluster analysis; Artificial neural network; Computer science; Artificial intelligence; Machine learning; Pattern recognition (psychology); Data mining","score_opus":0.02305575401566423,"score_gpt":0.24794935753964156,"score_spread":0.22489360352397733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2510065047","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23782592,0.00065695966,0.75792164,0.00019505536,0.000077916804,0.00007736887,0.00022963615,0.0006574188,0.0023580021],"genre_scores_gemma":[0.8783185,0.00023786623,0.11916593,0.000042326,0.000049523198,0.00009892457,0.00044170074,0.00004001613,0.0016052257],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985991,0.00042309336,0.00012416663,0.00033899024,0.00040887066,0.000105710715],"domain_scores_gemma":[0.99712616,0.0011257331,0.00044648105,0.00026510644,0.00095841975,0.00007815978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018098715,0.0008224357,0.0008781542,0.0021361734,0.00029662802,0.0007019168,0.0008845229,0.0007522961,0.0006123162],"category_scores_gemma":[0.0071002357,0.0003134383,0.0006983332,0.0017033395,0.00017087393,0.0012283184,0.00061688956,0.0006608397,0.00019468258],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016822608,0.00014903095,0.009934837,0.00008295691,0.00017584888,0.00007876406,0.000121275756,0.7916168,0.0026720937,0.0015292265,0.00101614,0.19245484],"study_design_scores_gemma":[0.0000026573707,0.000023967415,0.0024178072,0.0000069965135,0.000013777065,0.000015038973,0.000014159369,0.9956831,0.0010265318,0.00060379243,0.00018364181,0.000008505338],"about_ca_topic_score_codex":0.005142987,"about_ca_topic_score_gemma":0.005622838,"teacher_disagreement_score":0.005142987,"about_ca_system_score_codex":0.00071304204,"about_ca_system_score_gemma":0.00044843016,"threshold_uncertainty_score":0.0102261305},"labels":[],"label_agreement":null},{"id":"W2511803001","doi":"10.1145/2970276.2970326","title":"Deep learning code fragments for code clone detection","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":569,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Singapore Management University; University of Saskatchewan; National Science Foundation","keywords":"Computer science; Identifier; clone (Java method); Java; False positive paradox; Source code; Software maintenance; Code (set theory); Artificial intelligence; Software; Program comprehension; Point (geometry); Static program analysis; Software system; Programming language; Machine learning; Data mining; Software development; Set (abstract data type)","score_opus":0.019187734392043673,"score_gpt":0.27358695985009285,"score_spread":0.25439922545804916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2511803001","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19816624,0.0008791715,0.7896443,0.00042705837,0.00005540449,0.00015835722,0.00083335576,0.008314299,0.0015217531],"genre_scores_gemma":[0.636966,0.0002570652,0.3573798,0.00016544042,0.000039869577,0.00014471065,0.0028695604,0.0002719908,0.0019055692],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998789,0.00019553414,0.00009328692,0.0003758384,0.00043513157,0.00011128053],"domain_scores_gemma":[0.9940262,0.0023813006,0.0010413965,0.0009434537,0.0014115049,0.00019613923],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013687641,0.00095945573,0.00067886134,0.0031398525,0.00049929466,0.00079114275,0.0017591489,0.0011715156,0.0010724058],"category_scores_gemma":[0.009551055,0.000431241,0.00070280774,0.002056925,0.0006412547,0.0021297836,0.0011706733,0.0016725353,0.00064105034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029731842,0.00030751436,0.039417256,0.0002376953,0.0001153345,0.00028015685,0.0003551177,0.09719716,0.02623928,0.0031929598,0.006117604,0.8262427],"study_design_scores_gemma":[0.000013637969,0.000090111,0.004555907,0.00002722709,0.000033819513,0.00015403148,0.0000747449,0.97269565,0.013846747,0.006244814,0.0022469014,0.000016453432],"about_ca_topic_score_codex":0.0069904337,"about_ca_topic_score_gemma":0.009760031,"teacher_disagreement_score":0.0069904337,"about_ca_system_score_codex":0.0011084735,"about_ca_system_score_gemma":0.0012470287,"threshold_uncertainty_score":0.013899505},"labels":[],"label_agreement":null},{"id":"W2511811891","doi":"10.18293/seke2016-146","title":"Embedded Emotion-based Classification of Stack Overflow Questions Towards the Question Quality Prediction","year":2016,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Support vector machine; Stack (abstract data type); Perceptron; Quality (philosophy); Machine learning; Recall; Artificial intelligence; Multilayer perceptron; Software; Precision and recall; Ask price; Artificial neural network; Natural language processing; Programming language; Psychology; Cognitive psychology","score_opus":0.04203609170447039,"score_gpt":0.2968207305279166,"score_spread":0.2547846388234462,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2511811891","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93925464,0.00057631865,0.05473197,0.0004113817,0.00008856476,0.00020866748,0.0010600061,0.001035701,0.0026326862],"genre_scores_gemma":[0.97853506,0.00009415074,0.018401599,0.000055638906,0.000057420406,0.00009178591,0.0014432406,0.00003342538,0.0012877677],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979825,0.0007866676,0.00021405294,0.00040150664,0.00042901505,0.00018630901],"domain_scores_gemma":[0.9891263,0.0061084344,0.0012676157,0.00039545068,0.0027170493,0.00038513468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026156537,0.0007839429,0.00050656445,0.0024044742,0.00029889224,0.0011728962,0.00047184134,0.00092191703,0.0015083029],"category_scores_gemma":[0.013469923,0.00014856114,0.0006089946,0.0009371767,0.00030425686,0.0014530905,0.00082673493,0.00090697955,0.00096733565],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027968572,0.0010630839,0.3851142,0.00080809393,0.0002802834,0.00075275666,0.0045113093,0.020991707,0.07445049,0.0015633677,0.010173134,0.49749467],"study_design_scores_gemma":[0.000035308047,0.00041634738,0.24794163,0.000088468194,0.00013789753,0.00027421268,0.0015391227,0.72053087,0.023406586,0.001761058,0.003803176,0.0000652458],"about_ca_topic_score_codex":0.0017781139,"about_ca_topic_score_gemma":0.0016564225,"teacher_disagreement_score":0.0026156537,"about_ca_system_score_codex":0.0006218831,"about_ca_system_score_gemma":0.00029262167,"threshold_uncertainty_score":0.013833106},"labels":[],"label_agreement":null},{"id":"W2512560510","doi":"10.18293/seke2016-183","title":"The Software Architecture Mapping Framework for Managing Architectural Knowledge","year":2016,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Software architecture description; Architecture; Software architecture; Software engineering; Reference architecture; Architecture tradeoff analysis method; Resource-oriented architecture; Architectural pattern; Multilayered architecture; Software; Architecture framework; View model; Computer architecture; Software development; Component-based software engineering; Software construction; Programming language","score_opus":0.02189199255762845,"score_gpt":0.2568933819924493,"score_spread":0.23500138943482088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2512560510","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00081192405,0.00024308385,0.9890988,0.0007432278,0.000035250312,0.00025746835,0.00023889994,0.00093538465,0.0076359855],"genre_scores_gemma":[0.0183231,0.00046837082,0.977123,0.00015198787,0.000034576246,0.0005443367,0.00072610663,0.00016380903,0.0024646728],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9934329,0.003050696,0.0008816998,0.00070166105,0.0016040256,0.00032912454],"domain_scores_gemma":[0.99471647,0.0018874301,0.0005376109,0.0016181362,0.0008681313,0.00037222067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009479972,0.00168699,0.0009715173,0.011101761,0.0033820807,0.008681528,0.0035604362,0.0027645363,0.006910786],"category_scores_gemma":[0.01232365,0.0012917503,0.0039331163,0.009350412,0.0040099905,0.010934785,0.0069852783,0.0040889475,0.0021845833],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000016699363,0.0000670727,0.00083641434,0.0003052352,0.000083341845,0.00030459388,0.0020260694,0.010259448,0.0010133614,0.8582201,0.006368985,0.120498724],"study_design_scores_gemma":[0.000027325148,0.000051727868,0.0007064333,0.00048461964,0.00007252942,0.00070204365,0.0012492707,0.054427296,0.0015096959,0.7317334,0.20897506,0.00006069418],"about_ca_topic_score_codex":0.01703717,"about_ca_topic_score_gemma":0.01833665,"teacher_disagreement_score":0.01703717,"about_ca_system_score_codex":0.003769974,"about_ca_system_score_gemma":0.0070833066,"threshold_uncertainty_score":0.050135493},"labels":[],"label_agreement":null},{"id":"W2513125375","doi":"10.1016/j.jss.2014.09.042","title":"Cost, benefits and quality of software development documentation: A systematic mapping","year":2014,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":110,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Computer science; Quality (philosophy); Software engineering; Software; Systems engineering; Engineering; Operating system","score_opus":0.061711152586999704,"score_gpt":0.2973234557729851,"score_spread":0.23561230318598536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2513125375","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90123993,0.0804397,0.00841469,0.0023424136,0.00010015623,0.0011671906,0.0016485442,0.000070227194,0.004577251],"genre_scores_gemma":[0.9708166,0.014738273,0.013233881,0.00015195561,0.00004162595,0.00039202333,0.00036666944,0.00002466856,0.00023420797],"study_design_codex":"observational","study_design_gemma":"systematic_review","domain_scores_codex":[0.91834855,0.031553924,0.015609682,0.0025173714,0.031218413,0.00075208495],"domain_scores_gemma":[0.44498962,0.41771704,0.060907185,0.013112116,0.061730873,0.0015431235],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06531775,0.0006277206,0.001284803,0.028998695,0.0011401795,0.0033273643,0.0012409702,0.0011178332,0.0010226951],"category_scores_gemma":[0.29877087,0.00090120325,0.002827095,0.012455247,0.0019592189,0.007294583,0.0031152274,0.0013991578,0.00015113819],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013784342,0.00044910164,0.4879323,0.020741167,0.003196011,0.00026573546,0.00724907,0.002064102,0.0010954749,0.0030228607,0.001219202,0.47138652],"study_design_scores_gemma":[0.0004000324,0.0050650523,0.89552706,0.033781864,0.01247034,0.0020067464,0.014705259,0.008435883,0.0043945257,0.008323412,0.014577183,0.00031257913],"about_ca_topic_score_codex":0.004581434,"about_ca_topic_score_gemma":0.01534604,"teacher_disagreement_score":0.93468225,"about_ca_system_score_codex":0.0063546672,"about_ca_system_score_gemma":0.011005415,"threshold_uncertainty_score":0.3454374},"labels":[],"label_agreement":null},{"id":"W2514099307","doi":"10.1109/icws.2016.25","title":"What Do Client Developers Concern When Using Web APIs? An Empirical Study on Developer Forums and Stack Overflow","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; World Wide Web; Web service; Web API; Web application security; Web modeling; Mashup; Web development; Latent Dirichlet allocation; Popularity; Web application; Web standards; Topic model; Information retrieval","score_opus":0.0892087723956149,"score_gpt":0.3559306419840582,"score_spread":0.26672186958844335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2514099307","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99800223,0.00017295248,0.0005345895,0.00027950236,0.0000058370906,0.000040473027,0.00005879907,0.000010009926,0.00089561957],"genre_scores_gemma":[0.9983847,0.00020637961,0.00071419164,0.00008741396,0.000017060387,0.00009210996,0.00013952311,0.000014229901,0.00034438682],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9895536,0.0058029746,0.0009875523,0.00094217993,0.001908151,0.00080553495],"domain_scores_gemma":[0.71321625,0.23006126,0.03523239,0.0039784117,0.012896523,0.004615131],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.017416297,0.0004093058,0.00046939045,0.002943434,0.0020121655,0.0028918076,0.0008938689,0.001420447,0.0011900187],"category_scores_gemma":[0.11123859,0.0005074264,0.00042545446,0.0029315588,0.0017499261,0.0060104574,0.0021114582,0.0019502763,0.00030988466],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001956949,0.00039292613,0.80426246,0.00024652074,0.00004860034,0.00047003475,0.16121182,0.00022737817,0.0010569586,0.0008256978,0.0015350786,0.029526906],"study_design_scores_gemma":[0.000034272703,0.0002952716,0.7733263,0.00029261914,0.000067008055,0.00060254475,0.21035363,0.0053388085,0.0010106935,0.0007898939,0.007818086,0.00007089283],"about_ca_topic_score_codex":0.0051856725,"about_ca_topic_score_gemma":0.00707187,"teacher_disagreement_score":0.9825837,"about_ca_system_score_codex":0.0016546304,"about_ca_system_score_gemma":0.0014186511,"threshold_uncertainty_score":0.092107296},"labels":[],"label_agreement":null},{"id":"W2514591789","doi":"10.1007/s11334-016-0285-7","title":"Source code size prediction using use case metrics: an empirical comparison with use case points","year":2016,"lang":"en","type":"article","venue":"Innovations in Systems and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Computer science; Data mining; Source code; Metric (unit); Software metric; Machine learning; Univariate; Source lines of code; Software; Artificial intelligence; Multivariate statistics; Software quality; Software development; Engineering","score_opus":0.06943296041123372,"score_gpt":0.3112907257649176,"score_spread":0.24185776535368392,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2514591789","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9926917,0.00027134237,0.0046512433,0.000060490183,0.000014847697,0.00003176508,0.00062938937,0.00023013556,0.0014192634],"genre_scores_gemma":[0.9963085,0.000084846295,0.0020556173,0.000014561649,0.00001234985,0.00003310739,0.001040682,0.00006062223,0.00038968213],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9934676,0.0016290052,0.0005589505,0.00091303297,0.0031956544,0.00023569986],"domain_scores_gemma":[0.7835823,0.15775812,0.022105008,0.0104552535,0.023794685,0.002304568],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0063392045,0.0007743023,0.00051436643,0.007578823,0.00038185902,0.0013081976,0.0014054739,0.0011005921,0.0015377393],"category_scores_gemma":[0.10163636,0.00036191544,0.0007008688,0.0044334014,0.0006092056,0.0039565777,0.0011868085,0.0009642113,0.00079323765],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005936765,0.0005359396,0.92513674,0.00016464574,0.00019156099,0.00013151138,0.0007404209,0.004161819,0.0020488354,0.00035763293,0.0011564544,0.06478076],"study_design_scores_gemma":[0.000042126554,0.0010588752,0.8813798,0.00010158616,0.00022651517,0.00075305434,0.0008358282,0.10700489,0.0057435115,0.0010186804,0.0017795557,0.000055554003],"about_ca_topic_score_codex":0.002987244,"about_ca_topic_score_gemma":0.005136026,"teacher_disagreement_score":0.9936608,"about_ca_system_score_codex":0.00070954155,"about_ca_system_score_gemma":0.0005312701,"threshold_uncertainty_score":0.033525348},"labels":[],"label_agreement":null},{"id":"W2515047361","doi":"10.1109/bigmm.2016.36","title":"Empirical Investigation of Code and Process Metrics for Defect Prediction","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Eclipse; Machine learning; Software bug; Support vector machine; Artificial intelligence; Process (computing); Software metric; Component (thermodynamics); Random forest; Data mining; Software; Feature (linguistics); Binary classification; Code (set theory); Artificial neural network; Root cause; Software quality; Software development; Reliability engineering; Set (abstract data type); Programming language; Engineering","score_opus":0.05174375538336614,"score_gpt":0.3245794560113815,"score_spread":0.2728357006280154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2515047361","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9886757,0.00054249266,0.008886313,0.0001659449,0.000014092085,0.000018758983,0.000739942,0.000092154776,0.0008645495],"genre_scores_gemma":[0.9973436,0.000069216374,0.0016727856,0.000009304737,0.0000087060835,0.00001184357,0.0007476178,0.000012627907,0.00012437033],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99496907,0.002782382,0.0003128088,0.0005483116,0.0011624108,0.00022515244],"domain_scores_gemma":[0.86009806,0.11650538,0.010899478,0.006087464,0.0052628377,0.0011466857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0090325335,0.00062803744,0.00041570695,0.0035114065,0.0002796353,0.0007082404,0.00084972405,0.0007399907,0.00084256567],"category_scores_gemma":[0.07181337,0.00020218687,0.0006839569,0.0035142894,0.0005702067,0.001481209,0.00054125034,0.0011742646,0.00040800966],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012356478,0.00017189728,0.9553021,0.000048041187,0.00018402796,0.00008679528,0.00011804079,0.014838506,0.00024007447,0.00044616644,0.00077549816,0.02766518],"study_design_scores_gemma":[0.000014134153,0.00037584524,0.7205873,0.000051452927,0.00008976305,0.00039537504,0.00026772846,0.2744505,0.0009845583,0.0012256659,0.0015324946,0.000025132249],"about_ca_topic_score_codex":0.0046141776,"about_ca_topic_score_gemma":0.004277329,"teacher_disagreement_score":0.0090325335,"about_ca_system_score_codex":0.00042971122,"about_ca_system_score_gemma":0.00044364986,"threshold_uncertainty_score":0.04776919},"labels":[],"label_agreement":null},{"id":"W2516395483","doi":"10.18293/seke2016-150","title":"Efficiently Measuring an Accurate and Generalized Clone Detection Precision using Clone Clustering","year":2016,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Generality; Computer science; Measure (data warehouse); Software; Data mining; Variety (cybernetics); Cluster analysis; Java; Detector; Artificial intelligence; Machine learning; Algorithm; Programming language","score_opus":0.042422662228142616,"score_gpt":0.26906447764044567,"score_spread":0.22664181541230305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2516395483","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15993367,0.000796144,0.8312001,0.00018459081,0.00007773465,0.0001965154,0.000502468,0.0049429084,0.0021657976],"genre_scores_gemma":[0.52791494,0.00026256463,0.46915197,0.00013763583,0.000053972162,0.00019255301,0.00085522694,0.0004708692,0.0009602292],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9800597,0.0039847004,0.0017057416,0.0047737556,0.008716031,0.00076013623],"domain_scores_gemma":[0.9211428,0.03339469,0.01035474,0.018294418,0.016129736,0.00068366574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010428279,0.0014312208,0.0019710173,0.007477038,0.0011896426,0.0030796665,0.0024158244,0.0025517496,0.00050092424],"category_scores_gemma":[0.0660326,0.0007989468,0.0010779825,0.005889717,0.0013750669,0.0037153384,0.002578155,0.0016054912,0.000735968],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063385395,0.0003682876,0.13465692,0.0009256621,0.0009218876,0.00028315547,0.0017535953,0.10952673,0.10889558,0.008139665,0.0036442743,0.6302504],"study_design_scores_gemma":[0.00006491604,0.0005573414,0.074939065,0.0001791342,0.00034646777,0.00133606,0.00062025635,0.669384,0.22521667,0.019552676,0.0074225767,0.00038078253],"about_ca_topic_score_codex":0.0034703724,"about_ca_topic_score_gemma":0.0038197704,"teacher_disagreement_score":0.010428279,"about_ca_system_score_codex":0.0015405745,"about_ca_system_score_gemma":0.0015218903,"threshold_uncertainty_score":0.055150688},"labels":[],"label_agreement":null},{"id":"W2516609490","doi":"10.1145/2970276.2970362","title":"QUICKAR: automatic query reformulation for concept location using crowdsourced knowledge","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Information retrieval; Baseline (sea); Relevance (law); Software; Similarity (geometry); Data mining; Artificial intelligence; Programming language","score_opus":0.04795226298173493,"score_gpt":0.33766200119648654,"score_spread":0.2897097382147516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2516609490","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03132993,0.001200397,0.9091243,0.0011828499,0.00032029994,0.001595192,0.004503282,0.045815337,0.0049283886],"genre_scores_gemma":[0.14388537,0.00044558398,0.8398587,0.00054007367,0.00016449466,0.00081565423,0.0087771155,0.0014410209,0.0040719463],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99298054,0.0023539318,0.00052551006,0.0016445221,0.0021994617,0.00029614032],"domain_scores_gemma":[0.9883642,0.0063166865,0.00074545393,0.0022900333,0.0019457459,0.00033782455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004130936,0.00289341,0.0019723799,0.0061546657,0.0017755073,0.0022126345,0.0036096864,0.0023390015,0.01143079],"category_scores_gemma":[0.021858444,0.00080761936,0.0021146475,0.003473175,0.001352365,0.005590588,0.00535146,0.0026103337,0.0065826676],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012924536,0.0008183455,0.0031291791,0.0026103477,0.00028965162,0.0009845371,0.0060429093,0.015629077,0.0759433,0.011972392,0.080997825,0.8002899],"study_design_scores_gemma":[0.00066767045,0.0010615702,0.0038670625,0.0003386714,0.00039314266,0.0015700741,0.009618577,0.66606563,0.110451005,0.046089593,0.15938592,0.0004910892],"about_ca_topic_score_codex":0.010019377,"about_ca_topic_score_gemma":0.013418068,"teacher_disagreement_score":0.01143079,"about_ca_system_score_codex":0.0016618392,"about_ca_system_score_gemma":0.0035626395,"threshold_uncertainty_score":0.038239837},"labels":[],"label_agreement":null},{"id":"W2518214351","doi":"10.1145/2970276.2970348","title":"Migrating cascading style sheets to preprocessors by introducing mixins","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Maintainability; Style sheet; Cascading Style Sheets; Programming language; Code (set theory); Preprocessor; Code reuse; Semantics (computer science); Software engineering; World Wide Web; XML; Web page; Software","score_opus":0.009549300685594193,"score_gpt":0.2525183850635264,"score_spread":0.24296908437793222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2518214351","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1288091,0.00046664962,0.7736955,0.00024871324,0.00019133773,0.0005593658,0.00036891553,0.092085406,0.0035750093],"genre_scores_gemma":[0.2652029,0.00036655794,0.7111655,0.00039272144,0.00008636779,0.0002865823,0.0010074468,0.012916309,0.008575597],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99676275,0.00041929667,0.00045992003,0.00084475067,0.0013218623,0.00019135934],"domain_scores_gemma":[0.97975343,0.0059466683,0.0032659362,0.00818983,0.002307914,0.0005361465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024688202,0.0020466712,0.0006911245,0.0022257385,0.0006776035,0.0025226779,0.0015758781,0.0011033582,0.0024012134],"category_scores_gemma":[0.01620662,0.0013838606,0.0014160844,0.0012301456,0.0011011992,0.0035312003,0.0027306941,0.0019943556,0.002153126],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008838502,0.00033636246,0.039021757,0.00080648385,0.00025565166,0.0021847202,0.0035475935,0.012109632,0.23818795,0.011642654,0.012319466,0.67870384],"study_design_scores_gemma":[0.00009216191,0.00046893704,0.013029887,0.00030150157,0.00028991513,0.0027194882,0.0005784743,0.20850788,0.6364879,0.01700075,0.12019386,0.00032921127],"about_ca_topic_score_codex":0.0012108231,"about_ca_topic_score_gemma":0.0012305744,"teacher_disagreement_score":0.0025226779,"about_ca_system_score_codex":0.00077540136,"about_ca_system_score_gemma":0.0013430531,"threshold_uncertainty_score":0.013056576},"labels":[],"label_agreement":null},{"id":"W2520496598","doi":"10.1007/978-3-319-31545-4_14","title":"Monitoring and Controlling Release Readiness by Learning Across Projects","year":2016,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Psychology; Process management; Engineering","score_opus":0.021084402587033185,"score_gpt":0.27312677602538865,"score_spread":0.25204237343835545,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2520496598","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.170463,0.01059837,0.5974396,0.0026351684,0.00063259905,0.00025582415,0.00095708075,0.008663536,0.2083549],"genre_scores_gemma":[0.6205664,0.005407114,0.14798233,0.00034044348,0.00026204472,0.0002523756,0.0012713724,0.0016031102,0.2223148],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99909055,0.00012767073,0.00004265522,0.0002012974,0.00045644186,0.00008149067],"domain_scores_gemma":[0.9966018,0.0017510323,0.00043594246,0.00053242716,0.00048658805,0.00019217895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001381197,0.0008055825,0.000483409,0.00092202384,0.00031417818,0.0024124484,0.0012298017,0.0005916421,0.0073519587],"category_scores_gemma":[0.006227557,0.0003318954,0.00028444797,0.001049663,0.00048274687,0.0027416525,0.0011345699,0.0010213455,0.003358137],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008439283,0.00009320603,0.005889839,0.00015653571,0.0000146391585,0.00004308798,0.0005015079,0.0044015236,0.0071681077,0.009716461,0.00884811,0.9630827],"study_design_scores_gemma":[0.00007464465,0.0015229693,0.1566336,0.0013610977,0.00019933537,0.0011698223,0.0025996855,0.18085799,0.11424036,0.21297611,0.32796064,0.00040383433],"about_ca_topic_score_codex":0.0013256144,"about_ca_topic_score_gemma":0.0021538453,"teacher_disagreement_score":0.0073519587,"about_ca_system_score_codex":0.0006625749,"about_ca_system_score_gemma":0.00077928434,"threshold_uncertainty_score":0.024594784},"labels":[],"label_agreement":null},{"id":"W2520723151","doi":"10.7287/peerj.preprints.2373v1","title":"Stopping duplicate bug reports before they start with Continuous Querying for bug reports","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Data deduplication; Search engine indexing; Software bug; BitTorrent tracker; Security bug; Information retrieval; Software; Process (computing); Database; Programming language; Artificial intelligence; Cloud computing","score_opus":0.01488962674465569,"score_gpt":0.24715986460304953,"score_spread":0.23227023785839385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2520723151","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31673542,0.0051596374,0.6229573,0.002951275,0.0006053594,0.0017463678,0.0015559989,0.040889516,0.007399052],"genre_scores_gemma":[0.5562947,0.00083929865,0.43310174,0.0007578951,0.00030400485,0.00031397014,0.0021443078,0.0019985393,0.004245578],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.97844726,0.005896398,0.003069079,0.0033494625,0.008471,0.00076684507],"domain_scores_gemma":[0.80676454,0.10069859,0.026022708,0.042434957,0.021548502,0.0025306565],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021234063,0.001658433,0.0029882593,0.0059087775,0.0012342358,0.0047993283,0.003799966,0.0022436674,0.0025887326],"category_scores_gemma":[0.12853365,0.0009961155,0.0011025948,0.003675981,0.0016097212,0.0076442715,0.0036611976,0.002557806,0.0019871641],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014544276,0.00091457623,0.05702697,0.0019069434,0.00032405442,0.00064708886,0.0043186555,0.0054660067,0.051351078,0.00611398,0.013862133,0.8566141],"study_design_scores_gemma":[0.0011522331,0.008047968,0.13977969,0.0014460685,0.0016743157,0.009898221,0.007600933,0.3201671,0.35091066,0.03978027,0.11821634,0.0013261674],"about_ca_topic_score_codex":0.0018352949,"about_ca_topic_score_gemma":0.0018246147,"teacher_disagreement_score":0.021234063,"about_ca_system_score_codex":0.00093931094,"about_ca_system_score_gemma":0.0026234565,"threshold_uncertainty_score":0.11229783},"labels":[],"label_agreement":null},{"id":"W2523412570","doi":"10.1145/2961111.2962601","title":"Predicting Defectiveness of Software Patches","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code review; Computer science; Software quality; Context (archaeology); Code (set theory); Software engineering; Process (computing); Software; Reliability engineering; Programming language; Software development; Engineering; Geology","score_opus":0.015566090080215988,"score_gpt":0.24532930197120464,"score_spread":0.22976321189098867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2523412570","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98851657,0.001317151,0.006835773,0.00020151082,0.000048707338,0.00005366721,0.0018398839,0.00035126723,0.00083549734],"genre_scores_gemma":[0.9886113,0.00042561762,0.006182256,0.000047602247,0.00006575923,0.000026861435,0.0040649245,0.000047786416,0.0005278491],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99809366,0.0002433004,0.00022103332,0.0006879998,0.0006346558,0.00011929369],"domain_scores_gemma":[0.9791733,0.009880376,0.006304947,0.0012654456,0.0023918983,0.0009840053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014259422,0.0006497094,0.00058889046,0.0047614505,0.00031569804,0.0008514573,0.0006354198,0.0010216575,0.0015448587],"category_scores_gemma":[0.021423493,0.00022784335,0.0005239944,0.001842982,0.0005013014,0.0009456159,0.00062860694,0.0006280533,0.0008731646],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040654544,0.00027583478,0.9002639,0.00034898394,0.00016519615,0.00028602683,0.00018776246,0.014741857,0.0030783513,0.00036502065,0.0046057315,0.07527479],"study_design_scores_gemma":[0.000055067474,0.0008597785,0.79669666,0.00013369642,0.00023466065,0.0014809589,0.00033048666,0.18937597,0.005013645,0.002582395,0.003177095,0.000059653594],"about_ca_topic_score_codex":0.0039227284,"about_ca_topic_score_gemma":0.004130454,"teacher_disagreement_score":0.0047614505,"about_ca_system_score_codex":0.00045504537,"about_ca_system_score_gemma":0.0004769092,"threshold_uncertainty_score":0.007799804},"labels":[],"label_agreement":null},{"id":"W2525234853","doi":"","title":"Estimating the required test volume and effort for software verification and validation","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Software verification; Verification and validation; Volume (thermodynamics); Model validation; Software; Reliability engineering; Test (biology); Data mining; Software development; Statistics; Engineering; Software construction; Programming language; Mathematics; Data science","score_opus":0.021709473103936533,"score_gpt":0.276159024302778,"score_spread":0.2544495511988415,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2525234853","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11883012,0.000508349,0.87560064,0.00021098822,0.000029834704,0.00015376459,0.0003780569,0.0013444299,0.0029437486],"genre_scores_gemma":[0.6135191,0.0004940497,0.38191715,0.00007093783,0.00008246908,0.0004326937,0.0016595095,0.00042072043,0.0014034184],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9910642,0.0027783245,0.00056810636,0.0006077885,0.0045723203,0.00040924398],"domain_scores_gemma":[0.9444529,0.037278526,0.0058135404,0.004741924,0.0072074244,0.00050564343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046720514,0.0014475639,0.0011268986,0.0045788516,0.00036609016,0.0014997327,0.001782076,0.0011919354,0.0012820465],"category_scores_gemma":[0.05006033,0.00076650054,0.0010701928,0.0016067512,0.0004416557,0.0033507915,0.0010531422,0.0007789799,0.0006891221],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045607603,0.00028692296,0.034138385,0.00051972136,0.00019103615,0.00028026945,0.00024976252,0.5330913,0.034764882,0.010050616,0.0013433935,0.3846276],"study_design_scores_gemma":[0.00007284737,0.0006831897,0.025838178,0.00012089022,0.00013244085,0.00042668564,0.00013794871,0.9161986,0.041818947,0.010162169,0.004317223,0.00009088466],"about_ca_topic_score_codex":0.0023854384,"about_ca_topic_score_gemma":0.0040269373,"teacher_disagreement_score":0.0046720514,"about_ca_system_score_codex":0.0011718197,"about_ca_system_score_gemma":0.0015156355,"threshold_uncertainty_score":0.02470845},"labels":[],"label_agreement":null},{"id":"W2526402655","doi":"10.1109/jsyst.2015.2443049","title":"A Hybrid eBusiness Software Metrics Framework for Decision Making in Cloud Computing Environment","year":2015,"lang":"en","type":"article","venue":"IEEE Systems Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"National Natural Science Foundation of China","keywords":"Cloud computing; Computer science; Software; Software metric; Electronic business; Software development; Software quality; Business model; Business; Operating system","score_opus":0.05070912399453539,"score_gpt":0.3127849055961183,"score_spread":0.26207578160158296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2526402655","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022858065,0.00016748367,0.9963496,0.00016421243,0.000011967684,0.000067758505,0.000044458306,0.00024270554,0.00066594454],"genre_scores_gemma":[0.13508281,0.00029649434,0.86338377,0.00007025802,0.00003505279,0.00029343663,0.0002102205,0.00005326441,0.0005747462],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99579984,0.0016732095,0.00040882642,0.0004996753,0.0013644921,0.00025395918],"domain_scores_gemma":[0.99723804,0.0011869537,0.0003343782,0.00015862883,0.0009037795,0.00017823011],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0052477457,0.0017725223,0.0014099026,0.0037672317,0.0008817654,0.0030596503,0.0024076162,0.0012492797,0.0010872736],"category_scores_gemma":[0.007280681,0.00056442694,0.0012709536,0.0032104664,0.0008767006,0.0030272305,0.0018861791,0.0017553947,0.0003341578],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006534353,0.00013592497,0.0021527428,0.0002182251,0.0001244948,0.00020491672,0.00030712356,0.7090648,0.0020864801,0.103360265,0.0018037034,0.18047602],"study_design_scores_gemma":[0.0000071397308,0.000028961613,0.0002369451,0.00002591044,0.000013543428,0.000029003431,0.00003762318,0.9735474,0.00040447,0.024031779,0.0016207251,0.000016592963],"about_ca_topic_score_codex":0.013895027,"about_ca_topic_score_gemma":0.012468651,"teacher_disagreement_score":0.9947522,"about_ca_system_score_codex":0.002613028,"about_ca_system_score_gemma":0.0036795924,"threshold_uncertainty_score":0.027753055},"labels":[],"label_agreement":null},{"id":"W2526929964","doi":"10.1145/2976767.2976773","title":"The problems with eclipse modeling tools","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Eclipse; Computer science; Documentation; Model-driven architecture; Data science; Plug-in; World Wide Web; Software engineering; Unified Modeling Language; Programming language","score_opus":0.03583036983235569,"score_gpt":0.2436100963211499,"score_spread":0.20777972648879423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2526929964","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16538271,0.008916328,0.648357,0.07766242,0.0018504638,0.0008612682,0.004204876,0.016310357,0.07645451],"genre_scores_gemma":[0.44709316,0.003803858,0.49931043,0.00812412,0.0013149475,0.0013221691,0.00573618,0.011290738,0.022004364],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9148666,0.045211792,0.008232829,0.0056378236,0.024017904,0.0020330327],"domain_scores_gemma":[0.62438065,0.2700191,0.018615702,0.051056013,0.033334777,0.0025937215],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.098971955,0.0012313273,0.0012207435,0.0074131647,0.003506678,0.008418624,0.0030373442,0.0028523158,0.0032698184],"category_scores_gemma":[0.24703078,0.0026198893,0.001242844,0.007884289,0.0037127996,0.023550218,0.005304966,0.0062898872,0.0035359333],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005455451,0.00032470186,0.040681854,0.003032367,0.00021124825,0.001439236,0.07947736,0.0058997413,0.006826246,0.20981121,0.0954993,0.5562512],"study_design_scores_gemma":[0.00009214684,0.00018594356,0.015187319,0.0025047762,0.00010823837,0.002238639,0.0141617,0.024503974,0.004733733,0.09490762,0.8411418,0.00023422316],"about_ca_topic_score_codex":0.005146362,"about_ca_topic_score_gemma":0.00959645,"teacher_disagreement_score":0.098971955,"about_ca_system_score_codex":0.0030048285,"about_ca_system_score_gemma":0.005375749,"threshold_uncertainty_score":0.52342},"labels":[],"label_agreement":null},{"id":"W2527794739","doi":"","title":"Evolution and Architecture of Open Source Software Collections: A Case Study of Debian","year":2012,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Waterloo","keywords":"Architecture; Computer science; Software engineering; Computer architecture; Geography","score_opus":0.014184554057286978,"score_gpt":0.2357263345179535,"score_spread":0.22154178046066653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2527794739","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98957294,0.00027982076,0.001897204,0.00066041795,0.0000061365704,0.000068798356,0.00012663082,0.000098189994,0.0072898604],"genre_scores_gemma":[0.98314065,0.00035209424,0.008550462,0.00014392211,0.0000078492785,0.00006005932,0.00038788293,0.00011534794,0.0072417986],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.99783415,0.0007585855,0.000100871264,0.0003401695,0.00069111184,0.0002750734],"domain_scores_gemma":[0.9913277,0.003370671,0.0010850034,0.0010991968,0.0021389825,0.0009785078],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.0034446795,0.00026565167,0.0003305694,0.0034403773,0.004616859,0.0021495067,0.0015741423,0.00090323976,0.0012412322],"category_scores_gemma":[0.009230789,0.00039016138,0.00033776264,0.0037027118,0.0028143052,0.0036156145,0.0023915102,0.0010928212,0.00036638568],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004922379,0.0020925747,0.30858552,0.0004767363,0.00009780348,0.021774339,0.37830654,0.006976941,0.007760968,0.028868217,0.013520684,0.2310474],"study_design_scores_gemma":[0.000069977745,0.00079796737,0.45502886,0.00036975785,0.00009414883,0.00883662,0.23483269,0.024864653,0.008687359,0.009233127,0.2569308,0.00025414195],"about_ca_topic_score_codex":0.049374495,"about_ca_topic_score_gemma":0.100849494,"teacher_disagreement_score":0.99538314,"about_ca_system_score_codex":0.006264432,"about_ca_system_score_gemma":0.002558313,"threshold_uncertainty_score":0.098174214},"labels":[],"label_agreement":null},{"id":"W2528164061","doi":"","title":"Classification and retrieval of reusable object-oriented software designs","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Software construction; Software development; Data mining; Software; Software sizing; Software design; Software engineering; Information retrieval; Programming language","score_opus":0.03876485588082715,"score_gpt":0.27638572306459286,"score_spread":0.2376208671837657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2528164061","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19192438,0.001556957,0.79545295,0.00036926012,0.00006614867,0.0008627566,0.00077088113,0.0028991601,0.0060976483],"genre_scores_gemma":[0.27755412,0.0011001308,0.7130195,0.000097944285,0.00003134568,0.00032155728,0.003489385,0.00026969778,0.0041163433],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974222,0.0007258729,0.0003305912,0.0003123179,0.0010692768,0.00013979142],"domain_scores_gemma":[0.99378914,0.0015411236,0.00079858553,0.0018062997,0.0019147904,0.0001500123],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019551709,0.0005498097,0.00086586346,0.009144637,0.0008011276,0.0028198345,0.0010221569,0.0007094618,0.0009830236],"category_scores_gemma":[0.011497523,0.000346299,0.0012931278,0.005661755,0.00068286597,0.0024954134,0.0010999653,0.0004082359,0.0007353023],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016500102,0.00015386184,0.009566463,0.0006088636,0.000081427956,0.00035965405,0.0024923263,0.010673805,0.027041277,0.031783458,0.0046235295,0.91245025],"study_design_scores_gemma":[0.00021158699,0.0010553175,0.03644609,0.00081179535,0.00049453264,0.0035215698,0.004756888,0.5472791,0.09354199,0.1778058,0.13371311,0.00036222246],"about_ca_topic_score_codex":0.0047288565,"about_ca_topic_score_gemma":0.0039094407,"teacher_disagreement_score":0.009144637,"about_ca_system_score_codex":0.0013955309,"about_ca_system_score_gemma":0.0016543855,"threshold_uncertainty_score":0.010340035},"labels":[],"label_agreement":null},{"id":"W2529682206","doi":"10.1109/qrs.2016.21","title":"Model-Driven Evaluation of Software Architecture Quality Using Model Clone Detection","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software quality; Software architecture; Software engineering; Architecture; clone (Java method); Context (archaeology); Reference architecture; Quality (philosophy); Software; Architectural model; Software development; Programming language","score_opus":0.12407220115988676,"score_gpt":0.3610301952322685,"score_spread":0.23695799407238172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2529682206","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29259458,0.00016093806,0.7009575,0.0002637615,0.00003226069,0.00041076716,0.00015177095,0.0026256237,0.0028028253],"genre_scores_gemma":[0.71261483,0.000062462335,0.2860717,0.000049030987,0.00000796498,0.00024475664,0.0002730512,0.0002395047,0.0004367199],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98622745,0.0062638703,0.000816702,0.0007005531,0.0056574997,0.00033394134],"domain_scores_gemma":[0.9511866,0.0251611,0.0047342316,0.0068035885,0.011534625,0.0005797695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013997137,0.0011302065,0.00079052337,0.0040687886,0.0005639233,0.0028259675,0.0011921943,0.000974566,0.0010500932],"category_scores_gemma":[0.060454972,0.0005034918,0.0010840317,0.0018875154,0.00091426296,0.0031503742,0.0018057413,0.001194396,0.00019173801],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001010827,0.0013018835,0.09990486,0.0006467762,0.00075020274,0.00028970087,0.0027418467,0.43079358,0.07800027,0.032368403,0.0020535786,0.35013807],"study_design_scores_gemma":[0.000050461662,0.00047123645,0.009275331,0.000047684967,0.000074800024,0.0001035032,0.00026880734,0.953359,0.02909235,0.0062388577,0.0009492676,0.00006869485],"about_ca_topic_score_codex":0.004249362,"about_ca_topic_score_gemma":0.0050095785,"teacher_disagreement_score":0.013997137,"about_ca_system_score_codex":0.0024725643,"about_ca_system_score_gemma":0.0019505509,"threshold_uncertainty_score":0.0740248},"labels":[],"label_agreement":null},{"id":"W2529848334","doi":"","title":"Concept Vocabularies in Programmer Sociolects (Work in Progress)","year":2014,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Programmer; Computer science; Program comprehension; Identifier; Reading (process); Focus (optics); Affect (linguistics); Key (lock); Process (computing); Code (set theory); Software engineering; Programming language; Software; Human–computer interaction; World Wide Web; Software system; Linguistics; Operating system","score_opus":0.013917994797933873,"score_gpt":0.2503533044261394,"score_spread":0.23643530962820553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2529848334","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2891067,0.017389264,0.5330989,0.02570013,0.0017713776,0.00075722777,0.025496712,0.006491553,0.100188084],"genre_scores_gemma":[0.65900254,0.006837892,0.2737554,0.0017540639,0.0006608032,0.0005160052,0.031877816,0.0024729117,0.02312253],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9938466,0.0025378503,0.000704228,0.0012469867,0.0012865885,0.00037772013],"domain_scores_gemma":[0.9760969,0.0126253925,0.0011979615,0.0033938533,0.00537791,0.0013080409],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007905552,0.00052339316,0.00071104674,0.005474745,0.0020904874,0.0092977155,0.0014656989,0.001394239,0.01992668],"category_scores_gemma":[0.02688131,0.00070614985,0.0011152008,0.0057701706,0.0028642619,0.022406857,0.004024712,0.0028196832,0.0035908732],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003497883,0.00021329278,0.012915382,0.0012578187,0.00008014591,0.00017643649,0.01234897,0.002259701,0.005372361,0.5520781,0.035955742,0.37699223],"study_design_scores_gemma":[0.00017282148,0.00016489309,0.019602027,0.002166149,0.00013645265,0.0007115923,0.014360974,0.031129405,0.008383943,0.37604484,0.54685044,0.0002765414],"about_ca_topic_score_codex":0.020134427,"about_ca_topic_score_gemma":0.012193553,"teacher_disagreement_score":0.020134427,"about_ca_system_score_codex":0.003636733,"about_ca_system_score_gemma":0.00606456,"threshold_uncertainty_score":0.06666136},"labels":[],"label_agreement":null},{"id":"W2530596726","doi":"10.1109/tse.2017.2671865","title":"Measuring the Impact of Code Dependencies on Software Architecture Recovery Techniques","year":2017,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Ontario Ministry of Research, Innovation and Science; National Science Foundation","keywords":"Computer science; Granularity; Architecture; Software; Software architecture; Software architecture description; Implementation; Reference architecture; Graph; Code (set theory); Distributed computing; Software engineering; Theoretical computer science; Programming language","score_opus":0.027883247811874713,"score_gpt":0.27180466622482546,"score_spread":0.24392141841295076,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2530596726","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.885362,0.0014784344,0.10309223,0.00025207715,0.0000731609,0.00020415464,0.00062723685,0.0070555443,0.0018551663],"genre_scores_gemma":[0.8865195,0.00035654297,0.10997873,0.00006687454,0.000017731873,0.00012504263,0.0016071786,0.00062694174,0.0007014962],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9856695,0.0032600237,0.0012556417,0.0018969048,0.0071187983,0.00079923857],"domain_scores_gemma":[0.87981623,0.08321273,0.01118884,0.01572919,0.009383349,0.0006696154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005601871,0.001180549,0.00055388117,0.0028291296,0.0006103634,0.00076680304,0.0015153425,0.0009914893,0.0008156625],"category_scores_gemma":[0.090685554,0.00066809624,0.00080895994,0.00185802,0.0009452192,0.0028245621,0.0016106962,0.0018245269,0.00040646477],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013156764,0.0008209156,0.07999481,0.0012988121,0.00047861898,0.00042214216,0.0009820497,0.2815885,0.071679555,0.0018571776,0.0037370168,0.55582476],"study_design_scores_gemma":[0.00009122013,0.0020655296,0.08365214,0.00015554506,0.00041594292,0.00087531283,0.0005038338,0.749321,0.1544315,0.003006797,0.0053198743,0.00016139243],"about_ca_topic_score_codex":0.0029986552,"about_ca_topic_score_gemma":0.0041711554,"teacher_disagreement_score":0.005601871,"about_ca_system_score_codex":0.00081028923,"about_ca_system_score_gemma":0.0013798815,"threshold_uncertainty_score":0.029625893},"labels":[],"label_agreement":null},{"id":"W2530824252","doi":"10.1109/tse.2016.2616306","title":"A Framework for Evaluating the Results of the SZZ Approach for Identifying Bug-Introducing Changes","year":2016,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":190,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Queen's University; McGill University","funders":"","keywords":"Implementation; Computer science; Task (project management); Software bug; Software implementation; Software; Software engineering; Programming language; Systems engineering","score_opus":0.07051700283867052,"score_gpt":0.3266203963305439,"score_spread":0.2561033934918734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2530824252","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18175887,0.0011584981,0.7830005,0.0021481216,0.00022308582,0.0031583745,0.0072995215,0.010340092,0.010912854],"genre_scores_gemma":[0.41789708,0.00022154306,0.5741702,0.00030606065,0.0000618192,0.002099721,0.0040103984,0.00041689601,0.00081624935],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.91324455,0.034208298,0.009851528,0.007949765,0.032706924,0.0020389883],"domain_scores_gemma":[0.7617588,0.13147064,0.038196288,0.027574716,0.039200723,0.0017988258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08436929,0.0039090193,0.002666479,0.03423489,0.0021979169,0.007367328,0.0047058347,0.0032048132,0.0030984152],"category_scores_gemma":[0.23684873,0.0012255388,0.0031709776,0.012294005,0.0047069,0.008084452,0.0072175567,0.003259937,0.0012838282],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002032704,0.001332743,0.22531013,0.00316227,0.001773727,0.00048452243,0.006001778,0.13347536,0.020257903,0.05824849,0.018853698,0.52906674],"study_design_scores_gemma":[0.0008204114,0.0047197533,0.12419944,0.0008431581,0.000875353,0.0008194005,0.0039190142,0.7352008,0.033453107,0.07482974,0.019349,0.00097085745],"about_ca_topic_score_codex":0.010707685,"about_ca_topic_score_gemma":0.009767716,"teacher_disagreement_score":0.08436929,"about_ca_system_score_codex":0.003966857,"about_ca_system_score_gemma":0.004552765,"threshold_uncertainty_score":0.44619274},"labels":[],"label_agreement":null},{"id":"W2532475197","doi":"10.1109/wcre.2005.24","title":"Multiple Layer Clustering of Large Software Systems","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Cluster analysis; Computer science; Data mining; Software; Software system; Layer (electronics); CURE data clustering algorithm; Cluster (spacecraft); Correlation clustering; Artificial intelligence; Operating system","score_opus":0.017434554782714443,"score_gpt":0.2558569156586423,"score_spread":0.23842236087592789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2532475197","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1220902,0.00027025407,0.8726876,0.00011813114,0.00003641332,0.00016538963,0.00020001004,0.0026278691,0.0018041824],"genre_scores_gemma":[0.37717617,0.00022130943,0.61783034,0.000045632903,0.00001832241,0.00014615152,0.0009670299,0.0006302959,0.0029647679],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976981,0.00047101427,0.00021213155,0.00047150324,0.000888373,0.00025883003],"domain_scores_gemma":[0.99373853,0.0018153759,0.00078864716,0.0016628794,0.0017080376,0.00028657203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018148464,0.0009844064,0.00086671696,0.004056772,0.0014528843,0.0021514425,0.0012945037,0.00077515346,0.001810516],"category_scores_gemma":[0.009212769,0.00082108576,0.0011892851,0.0031524263,0.00073884946,0.0024746552,0.0021685173,0.0008567309,0.00095959747],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033958105,0.00014558308,0.016464952,0.00046189537,0.0003035607,0.00045698634,0.0016781416,0.6016975,0.040599786,0.0165273,0.0042190244,0.3171057],"study_design_scores_gemma":[0.000016259888,0.000060867253,0.0043158936,0.000030197782,0.000060914463,0.00015577751,0.00027943478,0.95829,0.01617182,0.016843485,0.0037274142,0.000047983238],"about_ca_topic_score_codex":0.0052925483,"about_ca_topic_score_gemma":0.008124451,"teacher_disagreement_score":0.0052925483,"about_ca_system_score_codex":0.0014742181,"about_ca_system_score_gemma":0.0014248176,"threshold_uncertainty_score":0.010696232},"labels":[],"label_agreement":null},{"id":"W2533383645","doi":"10.1002/smr.1819","title":"Error leakage and wasted time: sensitivity and effort analysis of a requirements consistency checking process","year":2016,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"National Science Foundation","keywords":"Consistency (knowledge bases); Computer science; Reliability engineering; Process (computing); Sequential consistency; Model checking; Consistency model; Data mining; Data consistency; Distributed computing; Engineering; Theoretical computer science; Programming language; Artificial intelligence","score_opus":0.018855989888243198,"score_gpt":0.2918275587189101,"score_spread":0.27297156883066687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2533383645","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9177408,0.00021068002,0.07844369,0.00036468473,0.000016533177,0.00019212862,0.0001995371,0.00036172217,0.0024702586],"genre_scores_gemma":[0.9909382,0.000025230545,0.008778078,0.00002427231,0.0000036475453,0.000039464932,0.000052235897,0.000019494033,0.00011942727],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.97561127,0.014653028,0.0010431329,0.0018159063,0.0058662184,0.0010104559],"domain_scores_gemma":[0.7015618,0.2592272,0.014036414,0.014878238,0.009272961,0.0010234878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019977346,0.00092098815,0.0008592595,0.0039742193,0.0006221279,0.0022948626,0.0013562377,0.0013164298,0.0012839792],"category_scores_gemma":[0.11953818,0.0007170827,0.0013832018,0.002519941,0.0020185614,0.002762468,0.0020037119,0.0017502233,0.00011745707],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009458681,0.00054805545,0.04111399,0.0002312593,0.00037119343,0.0004177798,0.0010564352,0.91074944,0.008701149,0.011975142,0.0002853997,0.023604287],"study_design_scores_gemma":[0.000024353154,0.00031552959,0.00734467,0.000030256557,0.00009499747,0.0000987711,0.0001921069,0.98087376,0.0057713008,0.0049890555,0.00022206211,0.000043229848],"about_ca_topic_score_codex":0.005773846,"about_ca_topic_score_gemma":0.0023242768,"teacher_disagreement_score":0.019977346,"about_ca_system_score_codex":0.0034181168,"about_ca_system_score_gemma":0.0014683973,"threshold_uncertainty_score":0.10565156},"labels":[],"label_agreement":null},{"id":"W2533477538","doi":"10.1109/wcre.2005.14","title":"Design Pattern Detection in Eiffel Systems","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Eiffel; Computer science; Software design pattern; Reverse engineering; Architectural pattern; Software engineering; Structural pattern; Engineering design process; Software; Software design; Software system; Programming language; Process (computing); Software development; Engineering; Object-oriented programming","score_opus":0.01870716233368852,"score_gpt":0.23124544406054579,"score_spread":0.21253828172685726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2533477538","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05107437,0.00024461627,0.92612165,0.0002659923,0.000064298285,0.00021428001,0.00063168374,0.019767497,0.0016155458],"genre_scores_gemma":[0.24162285,0.00022039084,0.75289905,0.00019854771,0.000016267897,0.00023694687,0.0014159367,0.0013137636,0.0020762985],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99102145,0.0029983,0.0012266614,0.0012484287,0.0030122634,0.0004928676],"domain_scores_gemma":[0.97511554,0.016128704,0.0029100785,0.0033681463,0.0022626084,0.00021505992],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005216712,0.0008508875,0.00068662607,0.0021913908,0.0009492171,0.0022351556,0.0016293477,0.0020501786,0.0021061671],"category_scores_gemma":[0.032116935,0.0016451456,0.0014322631,0.001412374,0.0013591608,0.003850122,0.0017796466,0.0021518255,0.0007395162],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012278109,0.0004789925,0.032533273,0.0032032905,0.00029219914,0.0072052325,0.006567406,0.13616744,0.08239911,0.08150276,0.013296906,0.6351257],"study_design_scores_gemma":[0.0003545341,0.000541427,0.0072523183,0.00077853963,0.00023144898,0.0050694398,0.0010663278,0.6142629,0.18588577,0.104132235,0.080149785,0.0002752559],"about_ca_topic_score_codex":0.0032443872,"about_ca_topic_score_gemma":0.0036583808,"teacher_disagreement_score":0.005216712,"about_ca_system_score_codex":0.0011062267,"about_ca_system_score_gemma":0.0015244605,"threshold_uncertainty_score":0.027588964},"labels":[],"label_agreement":null},{"id":"W2534105126","doi":"10.1109/step.2005.11","title":"Design Issues for Software Analysis and Maintenance Tools","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Software maintenance; Software engineering; Computer science; Software development; Software construction; Software system; Program comprehension; Social software engineering; Personal software process; Business process reengineering; Software development process; Package development process; Software peer review; Backporting; Software analytics; Software; Engineering; Programming language","score_opus":0.03676933256609241,"score_gpt":0.2968367665211923,"score_spread":0.2600674339550999,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2534105126","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0023093517,0.003991513,0.9667624,0.011250916,0.0005665941,0.0002911878,0.00004691607,0.0012083336,0.01357282],"genre_scores_gemma":[0.037649732,0.0021350759,0.9474429,0.0022181424,0.0007354233,0.0011002645,0.000112985574,0.00093953864,0.0076660938],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.95677775,0.020169334,0.0049452297,0.0038785932,0.012855183,0.0013738202],"domain_scores_gemma":[0.9112155,0.054139555,0.0039719124,0.013514416,0.015435051,0.0017236155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.037555743,0.002624487,0.0020790375,0.0041963416,0.0034484991,0.01840365,0.0060937833,0.010664161,0.00755076],"category_scores_gemma":[0.099327415,0.004503574,0.0025079565,0.0029156667,0.009423493,0.024859754,0.0053898054,0.012053972,0.0061732586],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007192466,0.00012073567,0.00044202208,0.0010041511,0.000078605415,0.0003315836,0.0024631098,0.0044736494,0.0033920982,0.87510747,0.010028405,0.10248634],"study_design_scores_gemma":[0.00010616359,0.00012427653,0.00019868994,0.000755736,0.000118964505,0.00060790364,0.0005224332,0.020100463,0.0033958845,0.7826193,0.19136284,0.00008728866],"about_ca_topic_score_codex":0.0010314825,"about_ca_topic_score_gemma":0.0008529908,"teacher_disagreement_score":0.037555743,"about_ca_system_score_codex":0.0032193826,"about_ca_system_score_gemma":0.004334203,"threshold_uncertainty_score":0.19861609},"labels":[],"label_agreement":null},{"id":"W2535082093","doi":"10.1109/wcre.2005.30","title":"Source versus Object Code Extraction for Recovering Software Architecture","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Source code; KPI-driven code analysis; Extractor; Object code; Code (set theory); Object (grammar); Architecture; Programming language; Software architecture; Software; Software engineering; Code generation; Software development; Operating system; Software quality; Artificial intelligence; Key (lock); Engineering","score_opus":0.018814147327450514,"score_gpt":0.275957905899446,"score_spread":0.2571437585719955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2535082093","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09067674,0.00074582035,0.8827345,0.00040155175,0.00007641251,0.00036247974,0.0012606982,0.020385284,0.003356566],"genre_scores_gemma":[0.1749707,0.00052682136,0.81552726,0.00011087159,0.00004121861,0.00013244458,0.0026958094,0.0025329494,0.0034620047],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9973456,0.000685222,0.0002666984,0.00034135912,0.0012224447,0.0001386724],"domain_scores_gemma":[0.9782738,0.010961308,0.0023363917,0.0050688633,0.0031697804,0.00018984143],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030379305,0.0011070252,0.000620995,0.004530487,0.00062660675,0.0010760793,0.00089554995,0.0011250334,0.0024372682],"category_scores_gemma":[0.02060204,0.0005953744,0.00087506865,0.0033385188,0.00077728456,0.0026325376,0.0012752205,0.0012046971,0.0019467477],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040988327,0.00020808553,0.015523128,0.00078303803,0.0001896488,0.0006313679,0.00093740085,0.007074834,0.07395987,0.00553961,0.006117791,0.88862544],"study_design_scores_gemma":[0.00023740601,0.00049121823,0.041124504,0.0003160845,0.00053201005,0.0026602403,0.00085639424,0.28815258,0.56912017,0.0237176,0.07253465,0.0002571666],"about_ca_topic_score_codex":0.0022275904,"about_ca_topic_score_gemma":0.004410187,"teacher_disagreement_score":0.004530487,"about_ca_system_score_codex":0.00037416906,"about_ca_system_score_gemma":0.0011166949,"threshold_uncertainty_score":0.016066313},"labels":[],"label_agreement":null},{"id":"W2536305519","doi":"10.1007/s10664-016-9462-4","title":"Multi-objective reverse engineering of variability-safe feature models based on code dependencies of system variants","year":2016,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Austrian Science Fund","keywords":"Reverse engineering; Exploit; Computer science; Software product line; Software; Feature engineering; Source code; Feature (linguistics); Software engineering; Data mining; Consolidation (business); Software system; Machine learning; Artificial intelligence; Software development; Programming language","score_opus":0.025584771873456095,"score_gpt":0.2587096846354994,"score_spread":0.23312491276204333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2536305519","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1924611,0.00009464167,0.80473703,0.00009928677,0.000017849296,0.000067673864,0.00006402566,0.00060577446,0.0018526273],"genre_scores_gemma":[0.8607123,0.000054679414,0.1379356,0.00002356348,0.0000063299035,0.000062434534,0.0001279091,0.00016915704,0.00090795726],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99823195,0.00059386337,0.00007056443,0.00023514299,0.0006829416,0.0001855678],"domain_scores_gemma":[0.9929703,0.004024583,0.00091594417,0.0011992027,0.00077362574,0.00011624894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002370166,0.00088155287,0.0007366521,0.000891601,0.00032842017,0.0009466706,0.0011145333,0.0006017129,0.0009833374],"category_scores_gemma":[0.008970257,0.00057716627,0.0013333204,0.00044854687,0.0008298545,0.0013426341,0.0011076228,0.0012542253,0.00013312441],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000065248285,0.00007602798,0.0020199227,0.00004675611,0.000048133006,0.000091655755,0.00006141118,0.9548068,0.005254285,0.008413845,0.00013062973,0.028985305],"study_design_scores_gemma":[0.0000034680595,0.000023274582,0.0001965889,0.0000031656555,0.000012103575,0.000011981802,0.000007789446,0.9947784,0.0013849495,0.003506627,0.000068508794,0.0000032809291],"about_ca_topic_score_codex":0.003411224,"about_ca_topic_score_gemma":0.0050999173,"teacher_disagreement_score":0.003411224,"about_ca_system_score_codex":0.00095046195,"about_ca_system_score_gemma":0.0013758268,"threshold_uncertainty_score":0.012534797},"labels":[],"label_agreement":null},{"id":"W2536386317","doi":"10.1145/2984043.2989224","title":"Removing stagnation from modern code review","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Eclipse; Software inspection; Code (set theory); Software engineering; Code review; Source lines of code; Process (computing); Software; Software development; Foundation (evidence); Programming language; Static program analysis; Software quality; Set (abstract data type)","score_opus":0.02755357530281521,"score_gpt":0.28104295670149054,"score_spread":0.25348938139867533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2536386317","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44600797,0.0060933763,0.5055732,0.0044208537,0.00057883945,0.0013999784,0.00074262323,0.020523585,0.014659487],"genre_scores_gemma":[0.6233943,0.0010562261,0.36505923,0.001093743,0.00025449853,0.0005445548,0.0014415559,0.000936885,0.0062189396],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98701364,0.005072922,0.0012498742,0.0024178002,0.003551719,0.0006940141],"domain_scores_gemma":[0.9260339,0.030062374,0.012429587,0.009886831,0.01824562,0.0033416362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014951065,0.0013542456,0.002188845,0.0038236028,0.0021233219,0.0030122043,0.0021600924,0.0016736032,0.0027598292],"category_scores_gemma":[0.07197867,0.0009770585,0.00090976653,0.0032858676,0.0010518883,0.004552962,0.00244556,0.0018413153,0.0022863518],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001165581,0.00086470693,0.055055376,0.0008889322,0.0002367711,0.00034598683,0.0018607009,0.04327701,0.0196099,0.008292424,0.035847943,0.8325547],"study_design_scores_gemma":[0.0008718534,0.003709579,0.058425948,0.00050724,0.00045758483,0.0019585455,0.0013628999,0.7741425,0.02905675,0.039081447,0.0900643,0.00036141693],"about_ca_topic_score_codex":0.005152503,"about_ca_topic_score_gemma":0.008512196,"teacher_disagreement_score":0.014951065,"about_ca_system_score_codex":0.0020502966,"about_ca_system_score_gemma":0.006636981,"threshold_uncertainty_score":0.07906973},"labels":[],"label_agreement":null},{"id":"W2536699905","doi":"10.1109/iccsii.2012.6454577","title":"On the effect of aspect-oriented refactoring on testability of classes: A case study","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Testability; AspectJ; Computer science; Programming language; Unit testing; Java; Aspect-oriented programming; Object-oriented programming; Cohesion (chemistry); Regression testing; Software engineering; Software; Software system; Reliability engineering; Software construction; Engineering","score_opus":0.027710526402926615,"score_gpt":0.3139488386688247,"score_spread":0.28623831226589813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2536699905","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99538773,0.00021522095,0.003130167,0.00011706361,0.0000047361414,0.00008633895,0.00005220421,0.000033449905,0.0009730711],"genre_scores_gemma":[0.99302316,0.00017197657,0.006438242,0.000042202144,0.000011164766,0.000070505645,0.00007395245,0.000021638532,0.00014714523],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98310375,0.010370738,0.0010439674,0.001204739,0.0035929875,0.00068387174],"domain_scores_gemma":[0.49255294,0.47720814,0.0118437745,0.009254144,0.008007841,0.001133227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013696871,0.0007444052,0.0006001178,0.0019230024,0.00077896623,0.00090095046,0.0013411478,0.0014042057,0.0008475747],"category_scores_gemma":[0.09103644,0.00032289638,0.0011682403,0.0015275192,0.0013322701,0.0012334972,0.00078045303,0.0012295605,0.00017793625],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039803474,0.018671136,0.4527704,0.0029454012,0.0013779951,0.012946298,0.013064665,0.11349133,0.055493128,0.0046064537,0.0012948948,0.3193579],"study_design_scores_gemma":[0.000987453,0.032185797,0.6676242,0.00071567873,0.0016296518,0.0061439425,0.0059500737,0.14611752,0.12830243,0.0038787886,0.0060635703,0.00040088894],"about_ca_topic_score_codex":0.003276321,"about_ca_topic_score_gemma":0.0034254,"teacher_disagreement_score":0.013696871,"about_ca_system_score_codex":0.0014082051,"about_ca_system_score_gemma":0.00083699066,"threshold_uncertainty_score":0.07243681},"labels":[],"label_agreement":null},{"id":"W2537446610","doi":"10.1145/3001867.3001874","title":"Towards predicting feature defects in software product lines","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software bug; Software quality assurance; Software product line; Software quality; Naive Bayes classifier; Machine learning; Classifier (UML); Artificial intelligence; Data mining; Quality assurance; Software; Context (archaeology); Software development; Support vector machine; Engineering","score_opus":0.015811707635195727,"score_gpt":0.2604466243986399,"score_spread":0.24463491676344418,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2537446610","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56159663,0.0009771909,0.4294337,0.0005464831,0.00006471646,0.00021002677,0.00092145795,0.004621309,0.0016284635],"genre_scores_gemma":[0.7864333,0.00024180385,0.21044664,0.00009150579,0.000036881356,0.00006262963,0.0014847001,0.00015388687,0.0010486805],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978114,0.0005805222,0.00014915258,0.00054471137,0.0007728019,0.00014149307],"domain_scores_gemma":[0.98018336,0.009997177,0.0026592154,0.0014156258,0.005382513,0.00036215264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033673486,0.0017873965,0.00080405024,0.005537635,0.00048734306,0.001813281,0.0011094655,0.0018398768,0.0007488025],"category_scores_gemma":[0.019784119,0.0005610516,0.0008177835,0.0017433795,0.00045074098,0.002346934,0.0006578712,0.0013587071,0.0010084597],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046182156,0.00088193204,0.28063846,0.0003094798,0.00019889833,0.00035356515,0.00038873588,0.29071575,0.017353417,0.0013354033,0.004625758,0.40273675],"study_design_scores_gemma":[0.000015210843,0.00014662267,0.012046093,0.000030837302,0.000033755903,0.00011027308,0.000076849006,0.97891104,0.0059866565,0.0019849278,0.00063629454,0.000021460119],"about_ca_topic_score_codex":0.010988384,"about_ca_topic_score_gemma":0.010483465,"teacher_disagreement_score":0.010988384,"about_ca_system_score_codex":0.00086483563,"about_ca_system_score_gemma":0.0008900702,"threshold_uncertainty_score":0.021848857},"labels":[],"label_agreement":null},{"id":"W2541131526","doi":"10.1109/step.2002.1267631","title":"Measurement and metrology requirements for empirical studies in software engineering","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software measurement; Metrology; Computer science; Software engineering; Software; Systems engineering; Software construction; Software development; Engineering; Programming language; Mathematics","score_opus":0.17829426053667294,"score_gpt":0.3820770247483785,"score_spread":0.20378276421170557,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2541131526","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060386225,0.013573968,0.84929603,0.06498815,0.0021172243,0.0027638234,0.0012219839,0.0004810951,0.059519086],"genre_scores_gemma":[0.19501424,0.008707501,0.7573344,0.01245576,0.0037361828,0.017517496,0.0012966981,0.0003298835,0.0036077597],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.714211,0.17126101,0.02742805,0.015206047,0.06938445,0.0025093933],"domain_scores_gemma":[0.28152,0.58505344,0.018933732,0.05902878,0.052703384,0.0027606618],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.24496843,0.0025606859,0.006396974,0.009524019,0.005017228,0.01308587,0.005651606,0.014658351,0.006802176],"category_scores_gemma":[0.58551574,0.0022638945,0.0042225276,0.013086437,0.033616513,0.043694478,0.015009558,0.020345323,0.0037011376],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034232708,0.000057928326,0.00064385764,0.0006070833,0.000038791335,0.000056304347,0.0009860286,0.00090707524,0.00018665534,0.9824881,0.0020216687,0.0119721275],"study_design_scores_gemma":[0.00009212009,0.00012074824,0.0015499924,0.0012444945,0.00003329141,0.00025878922,0.00084930874,0.0034420912,0.00018887977,0.96193105,0.030222831,0.00006642911],"about_ca_topic_score_codex":0.0031722274,"about_ca_topic_score_gemma":0.0017335436,"teacher_disagreement_score":0.7550316,"about_ca_system_score_codex":0.0072011133,"about_ca_system_score_gemma":0.015935026,"threshold_uncertainty_score":0.9310883},"labels":[],"label_agreement":null},{"id":"W2541219092","doi":"10.1109/ms.2016.156","title":"The Tragedy of Defect Prediction, Prince of Empirical Software Engineering Research","year":2016,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Software bug; Computer science; Field (mathematics); Tragedy (event); Empirical research; Software; Software engineering; Data science; Engineering; Programming language; Mathematics; Sociology","score_opus":0.043490389487501946,"score_gpt":0.32622678522319115,"score_spread":0.2827363957356892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2541219092","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.088315494,0.25946116,0.22329219,0.35747883,0.013961611,0.00036010527,0.008050736,0.0019140411,0.047165867],"genre_scores_gemma":[0.72216773,0.1164358,0.082453564,0.036765028,0.02365739,0.0005885656,0.0039692204,0.0009519453,0.013010861],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.93716365,0.021021396,0.0028648188,0.007950653,0.030388845,0.00061071815],"domain_scores_gemma":[0.40948546,0.4789419,0.031131428,0.041321024,0.035327643,0.0037924198],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.056613866,0.0014925586,0.0021008214,0.013391362,0.002703586,0.009565821,0.0030337365,0.0043927575,0.006284188],"category_scores_gemma":[0.36915067,0.0012006875,0.0010683622,0.015168,0.01302212,0.022393182,0.0047119176,0.0073725577,0.0030924676],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082785095,0.00041344023,0.092168555,0.0045706457,0.00089111924,0.00029924742,0.002014329,0.0102653755,0.0014919526,0.14727972,0.13178311,0.6079947],"study_design_scores_gemma":[0.000214524,0.0008291411,0.052688655,0.0056761974,0.00033299974,0.0012619786,0.0026148632,0.038321357,0.0029156485,0.6948282,0.19989395,0.00042248506],"about_ca_topic_score_codex":0.00417027,"about_ca_topic_score_gemma":0.003526899,"teacher_disagreement_score":0.94338614,"about_ca_system_score_codex":0.0025918693,"about_ca_system_score_gemma":0.0037297108,"threshold_uncertainty_score":0.2994063},"labels":[],"label_agreement":null},{"id":"W2541597427","doi":"10.1109/iccsii.2012.6454312","title":"On understanding software quality evolution from a defect perspective: A case study on an open source software system","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Software quality analyst; Software evolution; Computer science; Software quality; Software quality control; Software quality assurance; Quality assurance; Metric (unit); Quality (philosophy); Perspective (graphical); Software metric; Open source software; Software system; Software engineering; Software; Software maintenance; Open source; Software development; Software construction; Engineering; Artificial intelligence; Operations management; Operating system","score_opus":0.1127066081727112,"score_gpt":0.3738947011396486,"score_spread":0.2611880929669374,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2541597427","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99017626,0.00022744303,0.007942029,0.0004105947,0.0000031103104,0.00006219144,0.000057105248,0.000016566766,0.0011046713],"genre_scores_gemma":[0.98226315,0.00033644575,0.016701773,0.00006336513,0.000009946508,0.000036256875,0.00011063655,0.000015606085,0.00046291374],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9970868,0.0016112202,0.00017255373,0.0002962483,0.0006360754,0.00019709732],"domain_scores_gemma":[0.950956,0.041230157,0.0038020373,0.0011576017,0.002349913,0.0005043567],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004287522,0.0005080993,0.0004037043,0.0029118403,0.0011549862,0.001619874,0.0013027311,0.0021178015,0.00067806605],"category_scores_gemma":[0.022181993,0.000288607,0.00043688458,0.0034684269,0.0016096621,0.0034430888,0.0011432753,0.0011864298,0.00011506222],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064070866,0.0043838806,0.57311326,0.0015295768,0.00023735009,0.023606919,0.078098804,0.060821164,0.026100267,0.016156368,0.0018831205,0.21342869],"study_design_scores_gemma":[0.00015771585,0.0030647863,0.54794514,0.00050832226,0.00030242614,0.009612953,0.08411853,0.31398603,0.018282294,0.011364128,0.010425253,0.00023246935],"about_ca_topic_score_codex":0.011447366,"about_ca_topic_score_gemma":0.020064745,"teacher_disagreement_score":0.011447366,"about_ca_system_score_codex":0.0016292754,"about_ca_system_score_gemma":0.0007539231,"threshold_uncertainty_score":0.022761464},"labels":[],"label_agreement":null},{"id":"W2541904841","doi":"10.4230/lipics.ecoop.2016.8","title":"C++ const and Immutability: An Empirical Study of Writes-Through-const","year":2016,"lang":"en","type":"article","venue":"DROPS (Schloss Dagstuhl – Leibniz Center for Informatics)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Immutability; Computer science; Set (abstract data type); Benchmark (surveying); Code (set theory); Meaning (existential); State (computer science); Programming language; Epistemology; Philosophy; Blockchain; Computer security","score_opus":0.03031890135607083,"score_gpt":0.32196307523619433,"score_spread":0.2916441738801235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2541904841","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99444604,0.00026464966,0.002132677,0.00019797901,0.000011376043,0.00006519497,0.00040898882,0.00019805756,0.0022750602],"genre_scores_gemma":[0.9963995,0.00008364159,0.0017630131,0.000068134825,0.00000802612,0.000042135438,0.0008724261,0.00013285546,0.0006301361],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9842166,0.005806268,0.0016352049,0.0019129317,0.005611783,0.00081725745],"domain_scores_gemma":[0.67184716,0.25002286,0.0343194,0.02064986,0.0202631,0.002897563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012079672,0.00060319004,0.00035103946,0.0021130936,0.0012090419,0.0019856838,0.0016669426,0.0010187814,0.0017884431],"category_scores_gemma":[0.1641544,0.0004785887,0.000466869,0.0028001054,0.003158783,0.0043881787,0.0021996296,0.0027867625,0.0005807775],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006974049,0.00078839454,0.94122344,0.00048745872,0.00016184467,0.00065915973,0.008955702,0.0029320165,0.0026866505,0.0031661836,0.0031446465,0.035097048],"study_design_scores_gemma":[0.000089554625,0.0017342671,0.91007787,0.0003489431,0.00015828351,0.0026589443,0.01515234,0.0347479,0.0099279545,0.004199211,0.020784438,0.00012024832],"about_ca_topic_score_codex":0.005337648,"about_ca_topic_score_gemma":0.0060680388,"teacher_disagreement_score":0.012079672,"about_ca_system_score_codex":0.0009626674,"about_ca_system_score_gemma":0.0011626884,"threshold_uncertainty_score":0.0638842},"labels":[],"label_agreement":null},{"id":"W2543971965","doi":"10.1007/s10664-016-9452-6","title":"Review participation in modern code review","year":2016,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":116,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; McGill University","funders":"","keywords":"Software quality; Code review; Android (operating system); Open source software; Computer science; Process (computing); Source code; Open source; Software; Best practice; Quality (philosophy); Set (abstract data type); Data science; Software development; Political science; Operating system","score_opus":0.04737496011315095,"score_gpt":0.34805441343046845,"score_spread":0.3006794533173175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2543971965","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06788662,0.19137874,0.025517283,0.3875414,0.044408847,0.0045928066,0.0079069855,0.0018798423,0.26888746],"genre_scores_gemma":[0.58963215,0.08556865,0.019356158,0.10220861,0.038912006,0.011076773,0.008021806,0.0022755137,0.1429484],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.7297186,0.14960153,0.024533102,0.018250028,0.0680493,0.0098473495],"domain_scores_gemma":[0.2415008,0.40278572,0.073702544,0.054161023,0.18448032,0.043369595],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.16783057,0.0010418796,0.002366232,0.029018877,0.0050770035,0.013710721,0.0041694622,0.009296178,0.029567556],"category_scores_gemma":[0.62272364,0.0012152501,0.0014977419,0.01698516,0.003899182,0.012059562,0.014407983,0.004901427,0.008521236],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010248759,0.00015423223,0.024446221,0.015457166,0.0008522132,0.0007407256,0.018554412,0.00035596476,0.0027654185,0.028246896,0.5953298,0.31207213],"study_design_scores_gemma":[0.000110416535,0.00008550012,0.013006749,0.0053600217,0.00026244702,0.0003223674,0.0014473383,0.00029067337,0.00073126407,0.005720784,0.97260576,0.000056630928],"about_ca_topic_score_codex":0.0035625768,"about_ca_topic_score_gemma":0.009522077,"teacher_disagreement_score":0.8321694,"about_ca_system_score_codex":0.009304824,"about_ca_system_score_gemma":0.038137443,"threshold_uncertainty_score":0.8875835},"labels":[],"label_agreement":null},{"id":"W2544185825","doi":"10.1002/smr.1836","title":"Guest editor's introduction to the Special Issue on Program Comprehension (ICPC 2014)","year":2016,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Program comprehension; Computer science; Comprehension; Software engineering; Business process reengineering; Software; Inclusion (mineral); World Wide Web; Software system; Engineering; Operations management; Psychology; Programming language","score_opus":0.009207760569983105,"score_gpt":0.2733214875256557,"score_spread":0.2641137269556726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2544185825","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00020641975,0.008542905,0.0023710926,0.04971062,0.93138695,0.0000519228,0.00023585543,0.0005691553,0.006925011],"genre_scores_gemma":[0.0023790933,0.01691299,0.0027663342,0.029474596,0.88612664,0.00017553823,0.0006176095,0.0013387086,0.060208466],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.994296,0.00085404166,0.000643529,0.00082526176,0.002958337,0.00042279056],"domain_scores_gemma":[0.9665256,0.0067147133,0.0018840483,0.0011075442,0.019514214,0.0042538582],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006017196,0.0020825614,0.001986682,0.0048768963,0.0019563176,0.010039694,0.0025652924,0.004309613,0.0717523],"category_scores_gemma":[0.032133978,0.0006627969,0.0018739651,0.0032173542,0.0013806886,0.005914208,0.0032520955,0.007091464,0.03518318],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012419338,0.000012498542,0.00004219501,0.00017311788,0.0000058144465,0.000033245542,0.00003338489,0.000033412794,0.00009148198,0.0005036351,0.98255527,0.016503522],"study_design_scores_gemma":[0.0000079058855,0.000023042032,0.0001811959,0.00029088176,0.000009513491,0.00011883049,0.000054637712,0.00009395425,0.000109988316,0.00086469576,0.9982323,0.00001303659],"about_ca_topic_score_codex":0.00082271866,"about_ca_topic_score_gemma":0.001516376,"teacher_disagreement_score":0.0717523,"about_ca_system_score_codex":0.002562574,"about_ca_system_score_gemma":0.0039476166,"threshold_uncertainty_score":0.24003541},"labels":[],"label_agreement":null},{"id":"W2544306716","doi":"10.1109/dmesp.1991.171731","title":"Evaluating expert systems using a multiple-criteria, multiple-stakeholder approach","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Stakeholder; Computer science; Expert system; Field (mathematics); Term (time); Test (biology); Management science; Engineering; Artificial intelligence","score_opus":0.33405074992388245,"score_gpt":0.3647658618065172,"score_spread":0.03071511188263476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2544306716","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4049891,0.0004182574,0.5708455,0.00084662123,0.00007914618,0.0024072472,0.00019703557,0.0003807721,0.01983637],"genre_scores_gemma":[0.57046425,0.00015596359,0.42683798,0.000121998004,0.000025179506,0.00083280075,0.00018415757,0.000045320212,0.0013323483],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.92402893,0.055376757,0.003303931,0.0017779232,0.014840683,0.00067183026],"domain_scores_gemma":[0.831493,0.1204868,0.0063302387,0.006181376,0.03372306,0.0017855325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04857534,0.001348937,0.00085461995,0.005043052,0.0013741264,0.0033785114,0.0013789745,0.001623327,0.0022327758],"category_scores_gemma":[0.13447377,0.00033681083,0.00093757856,0.0023993317,0.0015525285,0.0027536154,0.002194029,0.0009914279,0.00038481414],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022709297,0.0027629843,0.048437987,0.003060119,0.001215151,0.000824375,0.014723836,0.1744345,0.03342252,0.038209304,0.0044317776,0.67620647],"study_design_scores_gemma":[0.0008773571,0.007996485,0.033508293,0.00083623605,0.0006405768,0.00085314823,0.011309161,0.8219893,0.04979939,0.05722675,0.014451476,0.00051175593],"about_ca_topic_score_codex":0.0019238171,"about_ca_topic_score_gemma":0.0038332993,"teacher_disagreement_score":0.04857534,"about_ca_system_score_codex":0.0024694602,"about_ca_system_score_gemma":0.002092442,"threshold_uncertainty_score":0.25689405},"labels":[],"label_agreement":null},{"id":"W2545010811","doi":"10.1007/s00766-016-0260-8","title":"Effective use of analysts’ effort in automated tracing","year":2016,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Lockheed Martin; National Aeronautics and Space Administration; National Science Foundation","keywords":"Tracing; Computer science; Programming language","score_opus":0.02469034825982708,"score_gpt":0.2809362799748374,"score_spread":0.2562459317150103,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2545010811","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6384915,0.001160962,0.33307543,0.0032562842,0.00015050129,0.0001833043,0.00008510828,0.0017472788,0.021849712],"genre_scores_gemma":[0.95523804,0.000131953,0.04336158,0.00014224833,0.00003276707,0.00004706778,0.0000370402,0.00012456512,0.00088483933],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.956258,0.02702475,0.0019742644,0.0021422263,0.01138952,0.0012111679],"domain_scores_gemma":[0.74789256,0.18272133,0.016927084,0.03307139,0.017244115,0.002143504],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024393635,0.0010898832,0.0007657617,0.004024087,0.0010718099,0.00454677,0.0017604898,0.0017365424,0.0012230747],"category_scores_gemma":[0.17848116,0.0007116368,0.00045351786,0.0018411374,0.001302697,0.0072318665,0.0028081015,0.0018502639,0.00044273178],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010005364,0.0012665618,0.04625039,0.00038483503,0.0003047119,0.00049945083,0.0064826068,0.042544838,0.020300148,0.025855538,0.0032281387,0.8518822],"study_design_scores_gemma":[0.00035437627,0.0021076738,0.055759456,0.00070364214,0.0008047039,0.0013422185,0.0059491387,0.7786018,0.055789586,0.07898416,0.019302333,0.0003009153],"about_ca_topic_score_codex":0.0026684161,"about_ca_topic_score_gemma":0.0041588736,"teacher_disagreement_score":0.024393635,"about_ca_system_score_codex":0.0015191075,"about_ca_system_score_gemma":0.004478144,"threshold_uncertainty_score":0.1290074},"labels":[],"label_agreement":null},{"id":"W2545762469","doi":"10.1002/smr.1821","title":"Detecting duplicate bug reports with software engineering domain knowledge","year":2016,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Domain (mathematical analysis); Software engineering; Domain knowledge; Context (archaeology); Data science; Software; Domain engineering; Software mining; Information retrieval; Data mining; Software development; Software construction","score_opus":0.008544446977499068,"score_gpt":0.24132643811420926,"score_spread":0.2327819911367102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2545762469","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7311422,0.0043116,0.24760741,0.0006735581,0.00016288264,0.0007268661,0.001155169,0.009061195,0.005159072],"genre_scores_gemma":[0.76084644,0.00052638113,0.23494263,0.00016760814,0.00008463968,0.00018651788,0.0016076248,0.0002041814,0.0014340722],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98883164,0.0030304475,0.00156454,0.0028016395,0.0034040425,0.0003675805],"domain_scores_gemma":[0.9287665,0.032685127,0.013965976,0.010529501,0.012960812,0.0010920252],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008593448,0.0012385342,0.001628596,0.01491687,0.001075093,0.0027198715,0.002029904,0.0015614823,0.0007390928],"category_scores_gemma":[0.06276806,0.000714826,0.001004281,0.006094849,0.0005774956,0.003758878,0.003936499,0.0011150127,0.0007506898],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004190094,0.000624444,0.1674639,0.0010169682,0.00048113414,0.00056080695,0.0020572094,0.012372885,0.017290896,0.00090834306,0.0039722007,0.7928323],"study_design_scores_gemma":[0.0002582663,0.0015470639,0.25982904,0.0007362017,0.0015162065,0.005270514,0.005499305,0.57692444,0.10788183,0.013624729,0.026340146,0.0005722816],"about_ca_topic_score_codex":0.004500883,"about_ca_topic_score_gemma":0.008025205,"teacher_disagreement_score":0.01491687,"about_ca_system_score_codex":0.00092407555,"about_ca_system_score_gemma":0.0020387268,"threshold_uncertainty_score":0.04544705},"labels":[],"label_agreement":null},{"id":"W2545852208","doi":"10.1109/dmesp.1991.171703","title":"Towards effective management of expert system projects","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Expert system; Computer science; Subject-matter expert; Domain (mathematical analysis); Software engineering; Development (topology); Systems engineering; Engineering; Artificial intelligence","score_opus":0.02376670100234926,"score_gpt":0.2626460874990867,"score_spread":0.23887938649673743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2545852208","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07101814,0.0021096058,0.82520175,0.02625885,0.0005411642,0.0042984234,0.00041662127,0.008859545,0.06129589],"genre_scores_gemma":[0.25775218,0.0017027694,0.7068309,0.001259915,0.00032946587,0.0039709066,0.0014837869,0.0005328419,0.02613724],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.93387836,0.03213072,0.0050680465,0.0049511315,0.020258129,0.0037135596],"domain_scores_gemma":[0.86063284,0.030467167,0.020916896,0.025862774,0.042397335,0.019723004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05965731,0.0011373141,0.0008948696,0.008665347,0.004275313,0.017003257,0.006803329,0.003847252,0.007558815],"category_scores_gemma":[0.14445941,0.0011108848,0.00048669218,0.005193217,0.002090033,0.01768629,0.011217444,0.0038001917,0.0050329487],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012981519,0.00079054903,0.008280289,0.00045227018,0.00004974486,0.00037094147,0.0054014917,0.008288368,0.0030534428,0.04730708,0.028544024,0.89733195],"study_design_scores_gemma":[0.000584146,0.001729614,0.03634235,0.0020857917,0.0001172832,0.0015421306,0.023578763,0.13089572,0.012298781,0.2527854,0.5376202,0.0004198735],"about_ca_topic_score_codex":0.0018689559,"about_ca_topic_score_gemma":0.0022379872,"teacher_disagreement_score":0.05965731,"about_ca_system_score_codex":0.0056765005,"about_ca_system_score_gemma":0.022262508,"threshold_uncertainty_score":0.31550175},"labels":[],"label_agreement":null},{"id":"W2547405428","doi":"10.1145/2950290.2950298","title":"API deprecation: a retrospective analysis and detection method for code examples on the web","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":68,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Application programming interface; Java; Code (set theory); Open source; World Wide Web; Programming language; Source code; Information retrieval; Software; Set (abstract data type)","score_opus":0.029438122143693974,"score_gpt":0.30323402722625226,"score_spread":0.2737959050825583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2547405428","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47215393,0.002405595,0.4486402,0.0019508507,0.00041292087,0.0030348527,0.025824247,0.019145323,0.026432103],"genre_scores_gemma":[0.4395987,0.000863238,0.52117246,0.00033418776,0.0001497697,0.0021014116,0.020580266,0.0018960191,0.013303872],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9913487,0.0014126656,0.0011129774,0.0020652493,0.0036701,0.00039025722],"domain_scores_gemma":[0.92099714,0.034719184,0.012962413,0.011133791,0.018888444,0.0012990684],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0081865825,0.00084249506,0.00055093557,0.016187502,0.001347395,0.0026340054,0.0018451167,0.0011708584,0.0033965784],"category_scores_gemma":[0.06695161,0.0006613841,0.00076702225,0.008543818,0.0011300758,0.0034219432,0.0024032863,0.0017219712,0.0031716244],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040344964,0.0003680208,0.27807665,0.0013914441,0.00017120368,0.0017269487,0.013988885,0.0020213753,0.015864579,0.013822669,0.040010955,0.63215387],"study_design_scores_gemma":[0.00011117652,0.000646461,0.33088413,0.0019260477,0.00058374123,0.008922549,0.015799206,0.17397432,0.068570934,0.023113498,0.37475792,0.00071003946],"about_ca_topic_score_codex":0.0070896116,"about_ca_topic_score_gemma":0.012575477,"teacher_disagreement_score":0.016187502,"about_ca_system_score_codex":0.00102041,"about_ca_system_score_gemma":0.0025543496,"threshold_uncertainty_score":0.043295264},"labels":[],"label_agreement":null},{"id":"W2547621596","doi":"10.1109/ms.2016.147","title":"Cyclomatic Complexity","year":2016,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":128,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Engineering and Physical Sciences Research Council","keywords":"Cyclomatic complexity; Computer science; Popularity; Metric (unit); Software metric; Software engineering; Code (set theory); Object (grammar); Software; Simple (philosophy); Software development; Programming language; Data science; Theoretical computer science; Software quality; Artificial intelligence; Engineering; Set (abstract data type); Political science; Epistemology","score_opus":0.03786635918510187,"score_gpt":0.27448962098181307,"score_spread":0.2366232617967112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2547621596","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07607019,0.01629217,0.52361053,0.015491977,0.00173112,0.0003795701,0.008064934,0.0012006299,0.35715887],"genre_scores_gemma":[0.78455913,0.011116377,0.14244206,0.003589384,0.0035892131,0.0008723651,0.0076759527,0.0006432466,0.045512315],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99352103,0.001075276,0.0005897947,0.0017089996,0.0025188138,0.0005860356],"domain_scores_gemma":[0.9868325,0.0061751823,0.0014109638,0.002218413,0.0025065101,0.0008564616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023766635,0.0010633502,0.0012210534,0.005402466,0.0022745803,0.00970993,0.0018716754,0.001781936,0.017023377],"category_scores_gemma":[0.015277977,0.00040400884,0.0014494971,0.00527762,0.005631787,0.010631867,0.0042184694,0.0033848383,0.0040753153],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000025251054,0.000016533264,0.00087077805,0.00014569386,0.000024940335,0.000053719003,0.00015916688,0.002891717,0.00029830704,0.9635796,0.010758803,0.021175455],"study_design_scores_gemma":[0.0000068710824,0.000017085373,0.0005952571,0.00005418893,0.0000115513385,0.00017299691,0.00007143899,0.0063550007,0.00037708678,0.94845784,0.043859255,0.00002143644],"about_ca_topic_score_codex":0.0019921428,"about_ca_topic_score_gemma":0.0011564724,"teacher_disagreement_score":0.017023377,"about_ca_system_score_codex":0.0046473746,"about_ca_system_score_gemma":0.0019301369,"threshold_uncertainty_score":0.0569489},"labels":[],"label_agreement":null},{"id":"W2548808595","doi":"10.1145/2993274.2993281","title":"Analysis of marketed versus not-marketed mobile app releases","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Software versioning; Duration (music); Minor (academic); Computer science; Mobile apps; App store; Software; World Wide Web; Operating system","score_opus":0.015973560002464134,"score_gpt":0.26833452855021966,"score_spread":0.2523609685477555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2548808595","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99370134,0.00038203745,0.0019654832,0.000064621665,0.000017416985,0.00009753152,0.0015308548,0.00014597925,0.0020947955],"genre_scores_gemma":[0.99158,0.00024344487,0.0031635482,0.000020889924,0.000019840014,0.00010795653,0.0033894293,0.000087752975,0.0013871428],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99676394,0.0004956319,0.00046773747,0.0005792667,0.0014695944,0.00022380863],"domain_scores_gemma":[0.9186032,0.053375598,0.01337963,0.0031189758,0.0098273065,0.0016952968],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035482666,0.0003229324,0.00031718294,0.0054323534,0.00036739424,0.0011514528,0.0004493362,0.000341228,0.0013590254],"category_scores_gemma":[0.033027846,0.00021503265,0.0006317388,0.0029456585,0.00047350823,0.0015257478,0.00065183145,0.00065173174,0.00030973903],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014940678,0.00052411214,0.7987951,0.0009076134,0.00044264703,0.0013945621,0.0048124967,0.00478176,0.02311623,0.00245705,0.002810347,0.15846407],"study_design_scores_gemma":[0.000012360409,0.00031195403,0.9798749,0.00005059167,0.00008729835,0.0003929252,0.0013407082,0.011477467,0.0025485153,0.0004582944,0.0034086304,0.000036322886],"about_ca_topic_score_codex":0.0023886613,"about_ca_topic_score_gemma":0.002760148,"teacher_disagreement_score":0.0054323534,"about_ca_system_score_codex":0.0005651917,"about_ca_system_score_gemma":0.00053865754,"threshold_uncertainty_score":0.018765211},"labels":[],"label_agreement":null},{"id":"W2550336446","doi":"10.1109/vlhcc.2016.7739671","title":"Developing usable APIs with XP and cognitive dimensions","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"USable; Usability; Computer science; Code refactoring; Application programming interface; Agile software development; Usability inspection; Documentation; Usability engineering; Software engineering; Cognitive walkthrough; Usability lab; Web usability; Pluralistic walkthrough; User interface; Human–computer interaction; Extreme programming; Process (computing); Usability goals; World Wide Web; Software development; Software; Software development process; Programming language","score_opus":0.023806520795622073,"score_gpt":0.26267398342973747,"score_spread":0.2388674626341154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2550336446","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25954655,0.0005098735,0.7045021,0.0011145682,0.000031756088,0.00082269363,0.00005316004,0.0008440526,0.032575253],"genre_scores_gemma":[0.28239182,0.00036006025,0.7129178,0.0001649517,0.0000143062425,0.001039729,0.00007408344,0.00013049389,0.0029066808],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9887983,0.0061024614,0.0007933786,0.0006783847,0.0032408582,0.00038654832],"domain_scores_gemma":[0.9704399,0.02035499,0.0015757929,0.0035490694,0.003433462,0.0006468028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012148921,0.00071438705,0.00029752555,0.0017368268,0.0009127354,0.0061038504,0.0012213844,0.00082329474,0.0013034318],"category_scores_gemma":[0.022879716,0.00060414313,0.0006673822,0.0012738197,0.0029365376,0.0052738045,0.004810983,0.001398567,0.0002617767],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022571275,0.0010461345,0.02473574,0.0017880218,0.00013993827,0.00082529534,0.05409274,0.0045216163,0.024030488,0.10678515,0.0027078362,0.7791013],"study_design_scores_gemma":[0.0006796337,0.0034851595,0.105687976,0.0033262128,0.00064911676,0.0042575337,0.06533328,0.108301714,0.10679712,0.35507277,0.24583231,0.0005772296],"about_ca_topic_score_codex":0.0015309303,"about_ca_topic_score_gemma":0.0015514023,"teacher_disagreement_score":0.012148921,"about_ca_system_score_codex":0.001344738,"about_ca_system_score_gemma":0.0028594395,"threshold_uncertainty_score":0.06425041},"labels":[],"label_agreement":null},{"id":"W2552403574","doi":"10.4204/eptcs.229.8","title":"Developing a Practical Reactive Synthesis Tool: Experience and Lessons Learned","year":2016,"lang":"en","type":"article","venue":"Electronic Proceedings in Theoretical Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Debugger; Key (lock); Computer science; Software engineering; Software; Process management; Debugging; Programming language; Engineering; Computer security","score_opus":0.029348351075375517,"score_gpt":0.32771308191577375,"score_spread":0.2983647308403982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2552403574","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22838953,0.001523032,0.74258244,0.0025329047,0.00020974608,0.00044870353,0.00013070526,0.004238279,0.019944528],"genre_scores_gemma":[0.4583415,0.0024921673,0.5184097,0.0010743367,0.00012373572,0.00032221296,0.00054702954,0.002942122,0.01574719],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.988701,0.0054460415,0.000689079,0.0014164322,0.0030268023,0.00072065287],"domain_scores_gemma":[0.9741896,0.016790286,0.0005132699,0.0026008987,0.004350766,0.0015551543],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017995974,0.0014615894,0.0008937491,0.0011479941,0.0018902115,0.0037778963,0.0037216733,0.0031599307,0.0056801867],"category_scores_gemma":[0.0429702,0.00093603175,0.00079441234,0.0010424819,0.0029043257,0.00789173,0.004742626,0.0035549407,0.00203277],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060191157,0.002420046,0.0089323195,0.0023710493,0.0001611451,0.004385538,0.09797183,0.028528066,0.071957394,0.018175764,0.010713934,0.7537811],"study_design_scores_gemma":[0.0005411645,0.006221555,0.0069063082,0.00219129,0.00043014623,0.016313896,0.046354137,0.1327262,0.15196007,0.03849512,0.59689295,0.00096716656],"about_ca_topic_score_codex":0.002074867,"about_ca_topic_score_gemma":0.002568452,"teacher_disagreement_score":0.017995974,"about_ca_system_score_codex":0.0012117393,"about_ca_system_score_gemma":0.0022212458,"threshold_uncertainty_score":0.09517294},"labels":[],"label_agreement":null},{"id":"W2553562410","doi":"10.1145/3193954.3193957","title":"UML diagram synthesis techniques","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Class diagram; Communication diagram; UML tool; Unified Modeling Language; Computer science; Applications of UML; Programming language; Activity diagram; Diagram; Context (archaeology); Process (computing); Use Case Diagram; Systems Modeling Language; System context diagram; Software; Database","score_opus":0.0157134826960976,"score_gpt":0.2774334927276555,"score_spread":0.26172001003155787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2553562410","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00024293305,0.0002200848,0.99448943,0.00008913455,0.00006399929,0.00013630156,0.00014103936,0.0018918401,0.0027252072],"genre_scores_gemma":[0.012954114,0.00068698137,0.9805304,0.00011957539,0.000059610258,0.00041942496,0.0006983746,0.0007482629,0.0037832346],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9941526,0.0019122983,0.00062747276,0.0010627466,0.002063765,0.0001810739],"domain_scores_gemma":[0.99123996,0.0049343444,0.0005282212,0.0014978141,0.0016908053,0.00010882528],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00442102,0.0017780149,0.00076293707,0.0047896313,0.0009847305,0.0023345582,0.0019298363,0.0016254267,0.021557054],"category_scores_gemma":[0.016285067,0.0012903295,0.0023650099,0.0023697093,0.0008837639,0.0028063809,0.0025356913,0.0020514892,0.01114652],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011690249,0.0001133635,0.0005958347,0.002605289,0.00020210001,0.000495648,0.001139487,0.020104915,0.032576434,0.27846813,0.021688066,0.64189386],"study_design_scores_gemma":[0.00023415397,0.00011829065,0.0003849273,0.0009997131,0.0002516258,0.0011353677,0.00030203722,0.11345265,0.07064009,0.22161305,0.59076196,0.00010610465],"about_ca_topic_score_codex":0.0011355446,"about_ca_topic_score_gemma":0.0013310804,"teacher_disagreement_score":0.021557054,"about_ca_system_score_codex":0.0009904258,"about_ca_system_score_gemma":0.0017679048,"threshold_uncertainty_score":0.07211554},"labels":[],"label_agreement":null},{"id":"W2555320227","doi":"","title":"Discussion on the Results of the Detection of Design Defects","year":2007,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal","funders":"","keywords":"Computer science; Code smell; Abstraction; Programming language; Precision and recall; Software engineering; Software; Software design; Vocabulary; Object-oriented design; Code (set theory); Software development; Software quality; Artificial intelligence","score_opus":0.025991823339252887,"score_gpt":0.24696537891214423,"score_spread":0.22097355557289133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2555320227","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28567383,0.020267097,0.41547158,0.074699424,0.004365578,0.00065107923,0.0043930938,0.004922909,0.1895554],"genre_scores_gemma":[0.8350473,0.0045565492,0.09031138,0.007779859,0.0013710754,0.00024768,0.0020904958,0.0010414781,0.057554204],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.976722,0.009453866,0.0007589185,0.0017202834,0.010284197,0.0010607195],"domain_scores_gemma":[0.8133614,0.14860304,0.0032821957,0.008973605,0.024613382,0.0011663954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022889841,0.0014368233,0.0011622271,0.0049629277,0.0011967905,0.003643069,0.0025937287,0.004869203,0.02339106],"category_scores_gemma":[0.09690704,0.00034985843,0.0026705982,0.0024645834,0.0019146074,0.002997206,0.0013068632,0.0017691618,0.0046901912],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006506301,0.0010849767,0.07916606,0.005610799,0.0008563042,0.0051069222,0.004397348,0.041000735,0.10386103,0.1173969,0.06968415,0.56532854],"study_design_scores_gemma":[0.0004925,0.0028518287,0.08567217,0.001855847,0.001632261,0.006396456,0.0043183076,0.07270444,0.5055981,0.10844405,0.20956214,0.0004718536],"about_ca_topic_score_codex":0.0028193374,"about_ca_topic_score_gemma":0.001317954,"teacher_disagreement_score":0.02339106,"about_ca_system_score_codex":0.0019095504,"about_ca_system_score_gemma":0.0010223653,"threshold_uncertainty_score":0.12105447},"labels":[],"label_agreement":null},{"id":"W2557778485","doi":"","title":"Code pattern analysis of object-oriented programming languages","year":2016,"lang":"en","type":"article","venue":"QSpace (Queen's University Library)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Queen's University","keywords":"Programming language; Computer science; Code (set theory); Object (grammar); Artificial intelligence; Set (abstract data type)","score_opus":0.006550633817031076,"score_gpt":0.20885508214265822,"score_spread":0.20230444832562713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2557778485","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79419345,0.0005800579,0.19737148,0.0003560356,0.000031709813,0.00034508418,0.0011106675,0.0019625707,0.0040489063],"genre_scores_gemma":[0.82963467,0.00034561643,0.16496885,0.00008480412,0.00001544554,0.00028801162,0.002050827,0.00032198685,0.002289729],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9974789,0.0004472432,0.00034775806,0.00040333657,0.0011454034,0.00017722962],"domain_scores_gemma":[0.98861605,0.0037877737,0.0028258662,0.0012000995,0.0033346652,0.00023552545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010978978,0.00035053233,0.00031822338,0.0052000233,0.00053146074,0.0010982084,0.0005389824,0.00037673593,0.0006206071],"category_scores_gemma":[0.010307259,0.00022307283,0.0005530893,0.0039696596,0.00049295236,0.0014063616,0.00074061746,0.00039693507,0.00022900352],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003529262,0.00030212404,0.46856037,0.0008174412,0.00018950125,0.0012527066,0.0037813012,0.008030932,0.044540744,0.007858254,0.0026597076,0.461654],"study_design_scores_gemma":[0.00008314456,0.00074260705,0.48902842,0.0004018916,0.0003370118,0.0055196392,0.0052068196,0.34187108,0.10322632,0.024831504,0.028586427,0.00016505478],"about_ca_topic_score_codex":0.0028275657,"about_ca_topic_score_gemma":0.0035922932,"teacher_disagreement_score":0.0052000233,"about_ca_system_score_codex":0.00050950074,"about_ca_system_score_gemma":0.0008823112,"threshold_uncertainty_score":0.005806327},"labels":[],"label_agreement":null},{"id":"W2558314336","doi":"10.15353/joci.v12i3.3280","title":"DataBasic: Design Principles, Tools and Activities for Data Literacy Learners","year":2016,"lang":"en","type":"article","venue":"The Journal of Community Informatics","topic":"Software Engineering Research","field":"Computer Science","cited_by":114,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"John S. and James L. Knight Foundation","keywords":"Computer science; Literacy; Set (abstract data type); Key (lock); Mathematics education; Information literacy; Learning design; Digital literacy; Pedagogy; Psychology; World Wide Web","score_opus":0.1930286443424438,"score_gpt":0.35142041559371073,"score_spread":0.15839177125126694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2558314336","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012434361,0.00042256748,0.972983,0.0015545264,0.0000439514,0.0008616267,0.00016010687,0.0032800252,0.008259871],"genre_scores_gemma":[0.026345545,0.00037382075,0.9675626,0.0002125397,0.000009810299,0.0010064277,0.00019332646,0.00042612775,0.0038698069],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99454224,0.0026008044,0.00079575845,0.00042706507,0.0013817558,0.00025234715],"domain_scores_gemma":[0.98976517,0.00516864,0.00068688305,0.001802264,0.0016895961,0.0008873863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012584842,0.001190237,0.0005868616,0.0029417889,0.0015466462,0.008280615,0.0027057943,0.00206503,0.0031260177],"category_scores_gemma":[0.016683647,0.001519216,0.00076706836,0.0013473495,0.0037555315,0.008113216,0.0049663344,0.002655867,0.0020957093],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002559716,0.0008104824,0.011053445,0.0036920155,0.00006536137,0.001041741,0.054041773,0.005965608,0.036195867,0.40305597,0.017156174,0.46666566],"study_design_scores_gemma":[0.00022600441,0.00054003135,0.0034443142,0.0022303876,0.00009878411,0.0031543868,0.009446988,0.022538776,0.02982107,0.1825277,0.74577487,0.00019669508],"about_ca_topic_score_codex":0.0011709817,"about_ca_topic_score_gemma":0.0021562295,"teacher_disagreement_score":0.012584842,"about_ca_system_score_codex":0.0011751138,"about_ca_system_score_gemma":0.004230268,"threshold_uncertainty_score":0.06655574},"labels":[],"label_agreement":null},{"id":"W2559532455","doi":"10.1007/s10664-016-9487-8","title":"Analysis of license inconsistency in large collections of open source projects","year":2016,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science","keywords":"License; Computer science; Source code; Header; MIT License; Computer security; Software; Open source; Reuse; World Wide Web; Open source software; Resource (disambiguation); Software engineering; Programming language; Engineering; Operating system","score_opus":0.032707972927695436,"score_gpt":0.309570352120259,"score_spread":0.27686237919256357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2559532455","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99428034,0.00041192118,0.0028627978,0.00011733776,0.000011102479,0.000043746513,0.0013195362,0.00014405536,0.00080909755],"genre_scores_gemma":[0.9850499,0.00020453318,0.0057945545,0.000039518287,0.00002654851,0.00009663141,0.008141516,0.00015377141,0.00049305265],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98369986,0.0051705604,0.0021011904,0.0023766495,0.005946326,0.00070544594],"domain_scores_gemma":[0.78367764,0.14514823,0.026481649,0.023157125,0.018857293,0.0026780665],"candidate_categories":["metaresearch","bibliometrics","open_science"],"consensus_categories":[],"category_scores_codex":[0.011193052,0.00044869655,0.00094146084,0.011125721,0.0023379826,0.003022758,0.0025525484,0.0015618518,0.0011659358],"category_scores_gemma":[0.12520209,0.0008879577,0.0008817926,0.013082663,0.0017320195,0.00411606,0.0035275128,0.0018818097,0.0004844033],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007523961,0.0010303247,0.8948003,0.0006214432,0.0007431725,0.0010888595,0.007189194,0.010667345,0.0044697467,0.0046228473,0.0055380394,0.06847631],"study_design_scores_gemma":[0.000115477385,0.00030992425,0.88815296,0.00020672653,0.00053077465,0.001848567,0.0073611843,0.0754903,0.006427916,0.008828933,0.010577682,0.00014962222],"about_ca_topic_score_codex":0.009044727,"about_ca_topic_score_gemma":0.010536984,"teacher_disagreement_score":0.99744743,"about_ca_system_score_codex":0.0016899647,"about_ca_system_score_gemma":0.002261096,"threshold_uncertainty_score":0.05919522},"labels":[],"label_agreement":null},{"id":"W2559873118","doi":"10.1109/issre.2016.12","title":"SV-AF — A Security Vulnerability Analysis Framework","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Vulnerability management; Computer science; Vulnerability (computing); Traceability; Computer security; Software security assurance; Scope (computer science); Secure coding; Security bug; Vulnerability assessment; Risk analysis (engineering); Software engineering; Information security; Business; Security service","score_opus":0.014440017623675798,"score_gpt":0.29040805345655735,"score_spread":0.27596803583288154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2559873118","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006829538,0.00033349884,0.98413384,0.00062190683,0.000030621366,0.00035268176,0.0010091614,0.00199503,0.0046937363],"genre_scores_gemma":[0.13121904,0.00041854766,0.8645033,0.000114516035,0.000055350574,0.00046849885,0.0018474482,0.00015097475,0.0012222561],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99333864,0.0025644933,0.000886092,0.00088139967,0.0019068478,0.00042253273],"domain_scores_gemma":[0.9869262,0.006186108,0.0021964563,0.0016218744,0.0026474718,0.0004218662],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009645338,0.001959218,0.0010758425,0.019314552,0.0018007613,0.0061704125,0.002746705,0.0019642576,0.0026591176],"category_scores_gemma":[0.01488677,0.0009016022,0.004084544,0.004919244,0.003506238,0.0062464452,0.0043400265,0.0023142297,0.00074490905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009603891,0.00022684838,0.027499488,0.00097134983,0.0006509417,0.0013660854,0.0028674204,0.12072754,0.0042710793,0.6571239,0.009845744,0.1743536],"study_design_scores_gemma":[0.0000271785,0.00013526849,0.0044887275,0.00065271836,0.00022573418,0.001538448,0.0016574574,0.5119966,0.0026049989,0.4348765,0.041661453,0.00013492245],"about_ca_topic_score_codex":0.013070705,"about_ca_topic_score_gemma":0.00792581,"teacher_disagreement_score":0.019314552,"about_ca_system_score_codex":0.0025298514,"about_ca_system_score_gemma":0.0055859745,"threshold_uncertainty_score":0.051010072},"labels":[],"label_agreement":null},{"id":"W2559885217","doi":"10.1007/s10664-017-9512-6","title":"Curating GitHub for engineered software projects","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":345,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Classifier (UML); Software engineering; Software development; Software bug; Software metric; Data mining; Machine learning; Data science; Software quality; Artificial intelligence; Programming language","score_opus":0.05729360034467611,"score_gpt":0.325474775312802,"score_spread":0.2681811749681259,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2559885217","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18679577,0.0020066497,0.5090892,0.0035381946,0.0012920207,0.0016227787,0.009279256,0.21141958,0.07495655],"genre_scores_gemma":[0.2968762,0.0013374245,0.58969104,0.00089080556,0.0002209033,0.00085102,0.023865888,0.048133526,0.038133223],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99223995,0.002226207,0.00042161884,0.0009854302,0.0035837567,0.0005430517],"domain_scores_gemma":[0.96813565,0.00811609,0.002003833,0.013489707,0.007088271,0.0011663334],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005722354,0.0015723514,0.00075341243,0.0062042875,0.0019351395,0.003147876,0.0018914911,0.0013301888,0.011715393],"category_scores_gemma":[0.056724243,0.0009064901,0.0015661491,0.0027672101,0.0010089177,0.004165006,0.008731061,0.0020568154,0.008287753],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006327671,0.00046045324,0.020304834,0.002400512,0.00022890615,0.002096733,0.005882422,0.005422005,0.030398984,0.024300568,0.18277103,0.7251008],"study_design_scores_gemma":[0.00029496904,0.00087832974,0.032284204,0.0022562635,0.0005327007,0.0038784377,0.005613033,0.1172522,0.08057525,0.08575327,0.67030215,0.00037927183],"about_ca_topic_score_codex":0.004172427,"about_ca_topic_score_gemma":0.010705533,"teacher_disagreement_score":0.011715393,"about_ca_system_score_codex":0.0009637127,"about_ca_system_score_gemma":0.0045226943,"threshold_uncertainty_score":0.039191842},"labels":[],"label_agreement":null},{"id":"W2560082121","doi":"10.1109/issre.2016.42","title":"Experience Report: An Empirical Study of API Failures in OpenStack Cloud Environments","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Cloud computing; Application programming interface; Computer science; Reliability (semiconductor); Operating system; Interface (matter); Database","score_opus":0.04239593228334955,"score_gpt":0.346668232704956,"score_spread":0.30427230042160647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2560082121","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99818027,0.00006030107,0.00056429586,0.00012371289,0.000003969848,0.00005840017,0.00021788414,0.000025279836,0.0007659646],"genre_scores_gemma":[0.9982634,0.00009679938,0.0007525824,0.00006927692,0.0000069895436,0.00009651372,0.0003447335,0.000027106471,0.00034264155],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9920035,0.003544847,0.0007421713,0.0007017649,0.0024348323,0.00057287305],"domain_scores_gemma":[0.8771788,0.07610286,0.021399593,0.004837394,0.016655525,0.0038257488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008253112,0.0004411867,0.00035925102,0.0023387757,0.0012947725,0.00146348,0.0013275747,0.00084896653,0.0011350942],"category_scores_gemma":[0.082162455,0.00046794701,0.00027970198,0.0021621624,0.0016895658,0.0034509231,0.0017051738,0.0016040999,0.0004261796],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026623454,0.001303792,0.8503666,0.0005341742,0.00009174237,0.0016432555,0.10918111,0.0009987743,0.0013307196,0.00071799743,0.0041124974,0.02945303],"study_design_scores_gemma":[0.0000493336,0.0018115514,0.7986845,0.00031559594,0.00009588637,0.0018217505,0.17366125,0.006866705,0.00242852,0.0007337014,0.013398202,0.00013296922],"about_ca_topic_score_codex":0.0067328927,"about_ca_topic_score_gemma":0.008216343,"teacher_disagreement_score":0.008253112,"about_ca_system_score_codex":0.0011110838,"about_ca_system_score_gemma":0.0011100464,"threshold_uncertainty_score":0.04364717},"labels":[],"label_agreement":null},{"id":"W2560290351","doi":"10.1109/mtd.2016.14","title":"Adjusting the Balance Sheet by Appending Technical Debt","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Technical debt; Debt; Obligation; Balance sheet; Liability; Recourse debt; Off-balance-sheet; Scope (computer science); Business; Internal debt; Debt levels and flows; Finance; Computer science; Software; Law; Political science; Software development","score_opus":0.013287639287752406,"score_gpt":0.25605461162922794,"score_spread":0.24276697234147554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2560290351","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.589077,0.011030814,0.19624932,0.008440983,0.0018888295,0.0004889286,0.0015018074,0.002627957,0.18869437],"genre_scores_gemma":[0.8984661,0.003952048,0.052405994,0.0009443914,0.00059742475,0.00019955629,0.0012542777,0.00050996937,0.04167017],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99690527,0.00068771496,0.0004964683,0.0004125894,0.0012425813,0.00025532913],"domain_scores_gemma":[0.98123837,0.005221037,0.0058130813,0.0033348212,0.0035180212,0.0008746139],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044135796,0.00087574223,0.0006303579,0.0025906418,0.0013310962,0.0057294425,0.0013589917,0.0021131882,0.011684169],"category_scores_gemma":[0.029899208,0.00046856093,0.00067653495,0.0028479549,0.0013183478,0.0079962285,0.0030754707,0.0019778362,0.0037248337],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008463801,0.00058966636,0.061163235,0.0012753203,0.0002563035,0.0037150003,0.0027722227,0.03563943,0.02211324,0.24739555,0.022106422,0.60212725],"study_design_scores_gemma":[0.00030898143,0.0011160155,0.08766102,0.0024404211,0.0005223894,0.004698937,0.0032778769,0.062282685,0.031590868,0.3772166,0.42838043,0.0005037991],"about_ca_topic_score_codex":0.0012013109,"about_ca_topic_score_gemma":0.0008135057,"teacher_disagreement_score":0.011684169,"about_ca_system_score_codex":0.0013200269,"about_ca_system_score_gemma":0.001454274,"threshold_uncertainty_score":0.039087474},"labels":[],"label_agreement":null},{"id":"W2560665552","doi":"10.1109/re.2016.61","title":"Requirements Engineering Visualization: A Systematic Literature Review","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Visualization; Computer science; Requirements engineering; Process (computing); Systematic review; Perspective (graphical); Software engineering; Data science; Requirements analysis; Data visualization; Information visualization; Software; Systems engineering; Data mining; Engineering; Artificial intelligence; Programming language","score_opus":0.01971300742911301,"score_gpt":0.2960467260133942,"score_spread":0.2763337185842812,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2560665552","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004818132,0.98168635,0.0037846528,0.0016112999,0.0002619478,0.0035664618,0.0024425243,0.00007122268,0.0017573996],"genre_scores_gemma":[0.033610545,0.94619006,0.010828234,0.0011098896,0.00010913161,0.0060559358,0.0016473212,0.000039702634,0.00040909214],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.974964,0.009452951,0.009596328,0.0013999443,0.004132241,0.00045443943],"domain_scores_gemma":[0.89923877,0.07310146,0.011486113,0.0022000275,0.013141294,0.0008324501],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026961993,0.0019239874,0.0055559915,0.047066495,0.0013108985,0.0028741825,0.0026712704,0.0022167251,0.004992478],"category_scores_gemma":[0.10682639,0.0014110631,0.0053406437,0.03378086,0.0012014227,0.0049807667,0.0031496044,0.001263158,0.00075305847],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010821774,0.00004232298,0.0010384455,0.86126965,0.0020296187,0.00043046926,0.0015235097,0.00034249428,0.00045519983,0.0009691944,0.0048746727,0.12691614],"study_design_scores_gemma":[0.00007759506,0.000120238685,0.0020415443,0.952192,0.007425493,0.0004976116,0.001590098,0.00016666175,0.00030276322,0.0007207664,0.03481562,0.000049564318],"about_ca_topic_score_codex":0.0069050547,"about_ca_topic_score_gemma":0.01934946,"teacher_disagreement_score":0.973038,"about_ca_system_score_codex":0.006085414,"about_ca_system_score_gemma":0.027639944,"threshold_uncertainty_score":0.1425904},"labels":[],"label_agreement":null},{"id":"W2560868365","doi":"10.1109/vissoft.2016.18keywords:","title":"Merge-Tree: Visualizing the Integration of Commits into Linux","year":2016,"lang":"en","type":"article","venue":"Software Visualization","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Commit; Computer science; Linux kernel; Operating system; Merge (version control); Programming language; Database; Parallel computing","score_opus":0.026638951978745665,"score_gpt":0.325634247231876,"score_spread":0.29899529525313034,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2560868365","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10224375,0.0017386475,0.67393863,0.0020093261,0.00059786875,0.00045584157,0.022373188,0.18035264,0.0162901],"genre_scores_gemma":[0.40191767,0.0015059817,0.5511207,0.00040293302,0.00013603496,0.00041608326,0.022020027,0.0154822515,0.006998345],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99953854,0.000093236544,0.000043470394,0.00008431079,0.00016532477,0.000075130745],"domain_scores_gemma":[0.9977095,0.00093513273,0.0002784834,0.00032395683,0.00046163317,0.00029131366],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010578647,0.0012353108,0.00046867863,0.0043722373,0.000810558,0.0020956087,0.0012798724,0.00086232467,0.00788889],"category_scores_gemma":[0.0057136575,0.00046810534,0.00084556924,0.0030208589,0.00044255034,0.002845988,0.0027067878,0.0016701627,0.0016566877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001980499,0.00050126074,0.046425164,0.0021881375,0.00031838714,0.0021922304,0.017935373,0.041383743,0.032440975,0.048255082,0.24564873,0.56073034],"study_design_scores_gemma":[0.000327999,0.00037636378,0.036361303,0.000700309,0.00023145574,0.002115788,0.005435446,0.42540392,0.04572005,0.07471425,0.40816063,0.00045253674],"about_ca_topic_score_codex":0.009147745,"about_ca_topic_score_gemma":0.011119317,"teacher_disagreement_score":0.009147745,"about_ca_system_score_codex":0.0005885953,"about_ca_system_score_gemma":0.0013982455,"threshold_uncertainty_score":0.02639097},"labels":[],"label_agreement":null},{"id":"W2561210185","doi":"10.1109/vissoft.2016.6","title":"Visualizing Project Evolution through Abstract Syntax Tree Analysis","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Commit; Computer science; Java; Programming language; Source code; Syntax; Source lines of code; Abstract syntax tree; Code (set theory); Software evolution; Software; Information retrieval; World Wide Web; Software engineering; Database; Software system; Parsing; Artificial intelligence; Set (abstract data type); Software construction","score_opus":0.03600958935723222,"score_gpt":0.3275286578723475,"score_spread":0.2915190685151153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2561210185","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23611493,0.0012616717,0.6680389,0.0013356343,0.00020766455,0.000239061,0.02335382,0.060592245,0.008856138],"genre_scores_gemma":[0.4697792,0.00095642736,0.5041574,0.0001577155,0.00006815799,0.0003477838,0.01565214,0.0055607078,0.0033204095],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9987457,0.0003538467,0.0001472496,0.00018758615,0.00047941014,0.00008611423],"domain_scores_gemma":[0.99070495,0.0038178212,0.0018111393,0.0009923247,0.002255834,0.00041790496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022007318,0.0007496746,0.0005232687,0.011477282,0.0005523428,0.0026954326,0.00073096994,0.00085861015,0.0030964552],"category_scores_gemma":[0.009109569,0.00046675306,0.00070284354,0.008291273,0.00036372544,0.0028527933,0.0017852827,0.0010792989,0.0010503908],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007701942,0.00035800863,0.08909488,0.0014210943,0.0003137662,0.0012444063,0.016129067,0.068904035,0.041691557,0.044421356,0.07695934,0.6586923],"study_design_scores_gemma":[0.00012894887,0.0002558072,0.07824588,0.0005010736,0.00023814266,0.0013034994,0.0041132667,0.65941226,0.034638766,0.07725713,0.14351915,0.00038592008],"about_ca_topic_score_codex":0.006806926,"about_ca_topic_score_gemma":0.0066654864,"teacher_disagreement_score":0.011477282,"about_ca_system_score_codex":0.0006534493,"about_ca_system_score_gemma":0.0012552239,"threshold_uncertainty_score":0.013534665},"labels":[],"label_agreement":null},{"id":"W2563491404","doi":"10.22215/etd/2012-07162","title":"Converting code clones to aspects using algorithmic approach","year":2012,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Library and Archives Canada","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Code (set theory); Programming language; Humanities; Philosophy","score_opus":0.038281742952649814,"score_gpt":0.3107449615388988,"score_spread":0.272463218586249,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2563491404","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018495627,0.000096723685,0.9684175,0.00015629276,0.00006410738,0.00022362384,0.000107514476,0.004195077,0.008243598],"genre_scores_gemma":[0.1258128,0.00031315474,0.8635978,0.00012105288,0.000037277932,0.0001991784,0.00082435564,0.0016892804,0.00740518],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980872,0.00031181876,0.00020494382,0.0004004783,0.00084053806,0.00015502151],"domain_scores_gemma":[0.99606967,0.0012989999,0.0002693366,0.0015332112,0.0007532331,0.000075454765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009979963,0.0006377621,0.00044775655,0.0017122225,0.0006797775,0.0030225324,0.0012544666,0.00088991417,0.006042889],"category_scores_gemma":[0.007501757,0.0007807323,0.0012834307,0.0014955243,0.0013856898,0.0021046475,0.002058439,0.0012190639,0.001887767],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018499467,0.0002658842,0.004788616,0.00062228483,0.000089030844,0.0010170997,0.0019484207,0.027223231,0.07253015,0.23867883,0.007507104,0.6451444],"study_design_scores_gemma":[0.00018009714,0.00036666586,0.0037366387,0.00046643615,0.00029721935,0.002955372,0.0011559427,0.29417008,0.15882502,0.2715656,0.26614633,0.00013454021],"about_ca_topic_score_codex":0.0012883517,"about_ca_topic_score_gemma":0.002274443,"teacher_disagreement_score":0.006042889,"about_ca_system_score_codex":0.000680834,"about_ca_system_score_gemma":0.0013358269,"threshold_uncertainty_score":0.020215511},"labels":[],"label_agreement":null},{"id":"W2563787900","doi":"","title":"Convertibility of function points to COSMIC-FPP: identification and analysis fo functional outliers","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Convertibility; Function point; Software; Computer science; Leverage (statistics); Function (biology); Identification (biology); Outlier; Data mining; Econometrics; Software development; Economics; Artificial intelligence","score_opus":0.013772303625680025,"score_gpt":0.2419740005070462,"score_spread":0.22820169688136618,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2563787900","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92998976,0.00024639862,0.060810745,0.000081111335,0.00002583409,0.00012628596,0.0027229544,0.0012275272,0.0047693043],"genre_scores_gemma":[0.9686903,0.00010071611,0.025002915,0.000011494371,0.000017072065,0.00015312388,0.0050595645,0.00016372614,0.000801022],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9923654,0.0010522234,0.0007247076,0.0008554454,0.0046312115,0.00037100835],"domain_scores_gemma":[0.9592144,0.019374412,0.0062766345,0.0067964206,0.007818306,0.00051990425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041880566,0.000936949,0.00091591524,0.011961797,0.00060762605,0.001697922,0.0010359657,0.0006238894,0.0013783252],"category_scores_gemma":[0.034755416,0.00025125337,0.00077690755,0.0113974195,0.0008708917,0.0016147447,0.0014830643,0.0009932417,0.00057315757],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074111787,0.00023691719,0.67630833,0.00026433877,0.00021634577,0.0010359967,0.0023842254,0.01770233,0.010211482,0.0075518256,0.0028478673,0.28049916],"study_design_scores_gemma":[0.000029819228,0.00053787854,0.83249587,0.00007830312,0.0000909799,0.0022435784,0.0020635135,0.115500085,0.026000636,0.008990709,0.011853112,0.00011546435],"about_ca_topic_score_codex":0.0027644786,"about_ca_topic_score_gemma":0.0015919071,"teacher_disagreement_score":0.011961797,"about_ca_system_score_codex":0.0008948428,"about_ca_system_score_gemma":0.00047977376,"threshold_uncertainty_score":0.022148848},"labels":[],"label_agreement":null},{"id":"W2569474846","doi":"10.4230/dagrep.6.4.80","title":"Natural Language Argumentation: Mining, Processing, and Reasoning over Textual Arguments (Dagstuhl Seminar 16161)","year":2016,"lang":"en","type":"article","venue":"DROPS (Schloss Dagstuhl – Leibniz Center for Informatics)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Argumentation theory; Argument (complex analysis); Natural (archaeology); Computer science; Cognition; Task (project management); Epistemology; Computational linguistics; Natural language; Linguistics; Natural language processing; Cognitive science; Psychology; Philosophy; Management; History","score_opus":0.008148196413969548,"score_gpt":0.2725896490723868,"score_spread":0.2644414526584173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2569474846","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07801486,0.100418925,0.3777217,0.22155565,0.040499818,0.0025124482,0.013858148,0.011473008,0.15394555],"genre_scores_gemma":[0.34815055,0.03991983,0.3045667,0.00874886,0.018643528,0.002905193,0.03259456,0.006249938,0.23822086],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98867583,0.0057170554,0.00069964764,0.0016874933,0.0023437976,0.00087624264],"domain_scores_gemma":[0.9921755,0.004062249,0.00031988675,0.0006357311,0.0011435479,0.0016630229],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017099842,0.0017611326,0.0016933114,0.0025643164,0.0021116433,0.0099743605,0.0022149445,0.0033127202,0.03109209],"category_scores_gemma":[0.023138598,0.0010712888,0.001885243,0.0023222112,0.0024641894,0.0056562466,0.007327095,0.005344171,0.024306022],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010684462,0.0009922883,0.0007843148,0.0011747703,0.00016265063,0.00041576676,0.0026790048,0.0052602547,0.005709956,0.08021108,0.51065826,0.39088327],"study_design_scores_gemma":[0.00038419766,0.0002444059,0.0035645447,0.0010454885,0.00005521217,0.0003725043,0.00077440793,0.00900627,0.006707051,0.18585454,0.7918742,0.000117183044],"about_ca_topic_score_codex":0.0011687563,"about_ca_topic_score_gemma":0.0011632738,"teacher_disagreement_score":0.03109209,"about_ca_system_score_codex":0.003730646,"about_ca_system_score_gemma":0.004761932,"threshold_uncertainty_score":0.10401338},"labels":[],"label_agreement":null},{"id":"W2570857834","doi":"10.1145/2990497","title":"Generating API Call Rules from Version History and Stack Overflow Posts","year":2017,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Android (operating system); Application programming interface; Cluster analysis; World Wide Web; Precision and recall; Set (abstract data type); Baseline (sea); Information retrieval; Operating system; Programming language; Machine learning","score_opus":0.08205151378663754,"score_gpt":0.30898722927561467,"score_spread":0.22693571548897712,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2570857834","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7576191,0.001642416,0.15724087,0.0008394619,0.00034930577,0.0014563078,0.039772183,0.0292039,0.011876453],"genre_scores_gemma":[0.7136769,0.00043713918,0.2050214,0.00030226167,0.00021290102,0.00077053724,0.07370201,0.00082483084,0.005052022],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949752,0.0006962207,0.0005335193,0.001818013,0.0016350652,0.00034204405],"domain_scores_gemma":[0.9783272,0.011727023,0.0024274092,0.0025618605,0.004424856,0.0005315783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037035209,0.0018536547,0.0010368619,0.010216343,0.0007798289,0.0019530877,0.0017545385,0.001707218,0.0012178429],"category_scores_gemma":[0.022524655,0.0006275139,0.0017708762,0.0038448845,0.000442361,0.0025217144,0.0011641434,0.0015111659,0.0021558683],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080456736,0.0013354184,0.37255377,0.0009625536,0.0006627453,0.0015676107,0.0014962233,0.027404282,0.012120432,0.002515914,0.0379222,0.5406543],"study_design_scores_gemma":[0.00013625043,0.000442441,0.16075072,0.00028575622,0.0005766766,0.0011099975,0.00073235534,0.77907217,0.02622111,0.0062276954,0.024246844,0.00019804499],"about_ca_topic_score_codex":0.014528703,"about_ca_topic_score_gemma":0.027987193,"teacher_disagreement_score":0.014528703,"about_ca_system_score_codex":0.0008861358,"about_ca_system_score_gemma":0.0021420554,"threshold_uncertainty_score":0.028888226},"labels":[],"label_agreement":null},{"id":"W2572410958","doi":"10.1109/icsme.2016.29","title":"Detecting Function Constructors in JavaScript","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Namespace; JavaScript; Computer science; Programming language; Unobtrusive JavaScript; Function (biology); Operating system; World Wide Web; Rich Internet application","score_opus":0.014089492151882752,"score_gpt":0.22694783489010145,"score_spread":0.2128583427382187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2572410958","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20386234,0.002051756,0.7170935,0.00021636697,0.00015692318,0.00026246064,0.002217814,0.06883706,0.0053017167],"genre_scores_gemma":[0.61629486,0.0009570312,0.36689225,0.0002683807,0.00007642377,0.00019964007,0.003853656,0.005629231,0.0058285603],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962702,0.000568142,0.000370688,0.0008658456,0.0017149189,0.00021022437],"domain_scores_gemma":[0.988996,0.005746236,0.001548514,0.0017068046,0.0017751431,0.00022722219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020393722,0.0012628014,0.0007062848,0.003748363,0.000722067,0.0017542404,0.0010953534,0.0012643599,0.0010682096],"category_scores_gemma":[0.010498524,0.000607018,0.0008808211,0.0014824232,0.0008404648,0.0024940786,0.0015395867,0.0010554233,0.0010156358],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078829436,0.00024206219,0.11155932,0.0013193084,0.00021513326,0.0016451284,0.0024616707,0.006702762,0.14543028,0.010379837,0.015934575,0.70332164],"study_design_scores_gemma":[0.00009318194,0.00032969768,0.08201184,0.0006484327,0.00030829164,0.0050998395,0.000789923,0.16132058,0.5973851,0.01617389,0.13550583,0.00033342568],"about_ca_topic_score_codex":0.0031393229,"about_ca_topic_score_gemma":0.00385326,"teacher_disagreement_score":0.003748363,"about_ca_system_score_codex":0.0007965215,"about_ca_system_score_gemma":0.0012661291,"threshold_uncertainty_score":0.010785341},"labels":[],"label_agreement":null},{"id":"W2574870096","doi":"10.1109/icsme.2016.45","title":"Continuous Maintenance","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"DevOps; Computer science; Agile software development; Troubleshooting; Server; Automation; Automatic summarization; Continuous production; Production (economics); Software engineering; Process (computing); Software; Engineering; World Wide Web; Operating system; Software deployment; Information retrieval","score_opus":0.01074600930452668,"score_gpt":0.23616580215158617,"score_spread":0.22541979284705949,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2574870096","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034898877,0.004120396,0.7678858,0.0038981002,0.0024887044,0.0015173841,0.0030595239,0.03478435,0.14734682],"genre_scores_gemma":[0.44876775,0.0033027583,0.402025,0.0025468685,0.0012592447,0.0018177241,0.009747013,0.0063166306,0.12421707],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9936405,0.0008870217,0.00054231,0.0012583133,0.0031757532,0.0004961378],"domain_scores_gemma":[0.96851146,0.005395288,0.0020709855,0.012060032,0.011056375,0.0009058223],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057589114,0.0014112714,0.0009327191,0.0027260934,0.001981569,0.0042837798,0.0046268185,0.0016176144,0.028646486],"category_scores_gemma":[0.028660089,0.00068962865,0.0011353537,0.0017552034,0.0014066483,0.0051301327,0.0038447548,0.001949907,0.012635135],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031785306,0.00025745362,0.0053552734,0.0012273304,0.00011857869,0.00058541656,0.001305449,0.008623482,0.011458414,0.062882,0.10959766,0.7982711],"study_design_scores_gemma":[0.0002041949,0.00063676224,0.0072609526,0.00071148574,0.00020583154,0.0027520817,0.00075699214,0.046193846,0.02009888,0.07420211,0.8467815,0.00019541226],"about_ca_topic_score_codex":0.0027346069,"about_ca_topic_score_gemma":0.0020788773,"teacher_disagreement_score":0.028646486,"about_ca_system_score_codex":0.0016070809,"about_ca_system_score_gemma":0.0033467521,"threshold_uncertainty_score":0.09583211},"labels":[],"label_agreement":null},{"id":"W2575124635","doi":"10.32920/ryerson.14654229","title":"Analysis on Relationship between Code Quality and Code Coverage in an XP Environment: a Case Study on the SWURV Project","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software quality; Code (set theory); Quality (philosophy); Code smell; Residual; Software; Code coverage; Extreme programming; Software engineering; Process (computing); Root cause; Static program analysis; Source code; Reliability engineering; Software development; Programming language; Software development process; Engineering; Algorithm; Set (abstract data type)","score_opus":0.2052220897337871,"score_gpt":0.4005590867465069,"score_spread":0.1953369970127198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2575124635","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99833894,0.000039642076,0.0011917925,0.000052873216,7.5903984e-7,0.000019939751,0.000032985386,0.00000884239,0.0003142903],"genre_scores_gemma":[0.99682957,0.00005793255,0.0027001973,0.000010268084,0.0000019025318,0.000021542883,0.00011872838,0.000009555704,0.00025029143],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9939824,0.0030776132,0.0002939432,0.0006868734,0.0016025148,0.00035671788],"domain_scores_gemma":[0.9064311,0.08040418,0.00478001,0.0016908474,0.0057941745,0.0008997282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005938465,0.00033259668,0.00037885198,0.002615379,0.0007270839,0.001013446,0.00078839064,0.00089799234,0.00063778757],"category_scores_gemma":[0.026331235,0.00027597471,0.0005903758,0.0019535015,0.0010157595,0.0009752408,0.0007838506,0.00095591036,0.000099291436],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031951966,0.0017855961,0.90095896,0.00029723404,0.00018519507,0.007065121,0.012835668,0.007832821,0.0042700614,0.0016744524,0.0006604385,0.062114883],"study_design_scores_gemma":[0.000054600623,0.0021139246,0.9074095,0.00012224252,0.00015752605,0.0034369435,0.014623149,0.063835,0.0054980507,0.0010851339,0.0016061567,0.000057803885],"about_ca_topic_score_codex":0.0070783356,"about_ca_topic_score_gemma":0.007991226,"teacher_disagreement_score":0.0070783356,"about_ca_system_score_codex":0.0010370486,"about_ca_system_score_gemma":0.0008317623,"threshold_uncertainty_score":0.031405985},"labels":[],"label_agreement":null},{"id":"W2577170234","doi":"10.7287/peerj.preprints.1920v1","title":"Analysis of Test Driven Development on sentiment and coding activities in GitHub repositories","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Java; Computer science; Coding (social sciences); Test-driven development; Set (abstract data type); Software; Control (management); Productivity; World Wide Web; Software engineering; Software development; Operating system; Programming language; Artificial intelligence; Mathematics; Statistics","score_opus":0.01393982838113878,"score_gpt":0.2520172469555547,"score_spread":0.23807741857441594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2577170234","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9979366,0.00006755531,0.00077151216,0.000065487344,0.0000054183615,0.000021109507,0.0003518419,0.000037941893,0.0007424087],"genre_scores_gemma":[0.99753296,0.00006140495,0.0010462556,0.00002477974,0.000013579731,0.000030363117,0.0009406491,0.000020463503,0.000329572],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.995141,0.0015623778,0.00044110912,0.0003477778,0.0021276043,0.00038011788],"domain_scores_gemma":[0.9426059,0.029916571,0.014031223,0.0017336966,0.010563905,0.0011487714],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038954497,0.00029058108,0.00041702206,0.0045439736,0.0003417825,0.0012507823,0.00036686365,0.0002962013,0.0005142069],"category_scores_gemma":[0.03528881,0.00016523332,0.00034161945,0.0040374044,0.00040011166,0.0010092335,0.0011076444,0.00040011606,0.00018635526],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053929136,0.00022202065,0.8990643,0.0002903823,0.00012220684,0.0004538293,0.0049845204,0.00094826886,0.010623547,0.00030613586,0.0015321318,0.08091345],"study_design_scores_gemma":[0.0000074334143,0.0001754004,0.98653525,0.000042990316,0.000036634392,0.00027600434,0.0027545767,0.005568785,0.003020413,0.00014330447,0.0014109997,0.00002813293],"about_ca_topic_score_codex":0.002857902,"about_ca_topic_score_gemma":0.0027216699,"teacher_disagreement_score":0.0045439736,"about_ca_system_score_codex":0.0008741361,"about_ca_system_score_gemma":0.00047826464,"threshold_uncertainty_score":0.020601332},"labels":[],"label_agreement":null},{"id":"W2578208870","doi":"10.1109/icsme.2016.62","title":"BigCloneEval: A Clone Detection Tool Evaluation Framework with BigCloneBench","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":82,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Java; Computer science; Benchmark (surveying); Precision and recall; Open source; Software; Software engineering; Variety (cybernetics); Operating system; Artificial intelligence; Biology; Gene","score_opus":0.01987006014401904,"score_gpt":0.27682893164683936,"score_spread":0.2569588715028203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2578208870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.097125754,0.002003185,0.4833073,0.00046420394,0.00023728535,0.0020780007,0.013997594,0.39380452,0.006982217],"genre_scores_gemma":[0.23349842,0.00052695296,0.69287187,0.000617405,0.00010974209,0.00298646,0.043178443,0.023098033,0.0031126258],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9818018,0.004794022,0.0027442912,0.0026137854,0.0073217046,0.00072451215],"domain_scores_gemma":[0.9526215,0.024881823,0.004464408,0.0075087384,0.009325706,0.0011977932],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013150207,0.003879452,0.001755994,0.014048693,0.0013167296,0.004696232,0.005124886,0.002830743,0.0035039615],"category_scores_gemma":[0.055462334,0.0015012143,0.0025904581,0.0062465705,0.0017311275,0.007081707,0.0048633213,0.0022228158,0.0024214147],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028415604,0.0022817592,0.08306343,0.005242383,0.0019632285,0.001296721,0.00400197,0.04574897,0.058216218,0.018089268,0.15064046,0.6266141],"study_design_scores_gemma":[0.0014280095,0.004016141,0.043023434,0.0009571306,0.0009477713,0.002528487,0.0015961407,0.5916052,0.19737802,0.024822416,0.1306329,0.0010644094],"about_ca_topic_score_codex":0.0073636053,"about_ca_topic_score_gemma":0.0069713113,"teacher_disagreement_score":0.9868498,"about_ca_system_score_codex":0.0022221801,"about_ca_system_score_gemma":0.0030595819,"threshold_uncertainty_score":0.069545746},"labels":[],"label_agreement":null},{"id":"W2579161546","doi":"10.1109/tse.2017.2654244","title":"Using Natural Language Processing to Automatically Detect Self-Admitted Technical Debt","year":2017,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":220,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Technical debt; Computer science; Debt; Quality (philosophy); Bad debt; Code (set theory); Finance; Software development; Business; Software; Programming language; Set (abstract data type)","score_opus":0.016474308440274747,"score_gpt":0.29079089578154826,"score_spread":0.2743165873412735,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2579161546","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54269034,0.0011948379,0.3868244,0.0026460127,0.000270334,0.00097418454,0.016264506,0.04292936,0.006205924],"genre_scores_gemma":[0.51156276,0.00048842287,0.44539315,0.000763437,0.00014629724,0.0008269796,0.036552895,0.0009254484,0.0033406785],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954224,0.0012123146,0.0006210185,0.00096207304,0.0016246103,0.00015760571],"domain_scores_gemma":[0.95295626,0.028992673,0.008670548,0.0023157916,0.006552402,0.0005124075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033610878,0.0011235254,0.0005177744,0.0055928575,0.0006546698,0.001574574,0.0014790483,0.0013845349,0.0010028395],"category_scores_gemma":[0.0233499,0.0005565571,0.000801995,0.0028632693,0.00073794526,0.0033241936,0.001387197,0.0016021637,0.0010208106],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009374428,0.0010531974,0.11374132,0.004000246,0.00021022647,0.005476706,0.006697209,0.014578611,0.16033259,0.0069390493,0.05151874,0.6345147],"study_design_scores_gemma":[0.00025091754,0.00061347574,0.08801452,0.00054865493,0.00023020108,0.0036057723,0.0041336655,0.73614323,0.07241608,0.03161733,0.06215944,0.00026674054],"about_ca_topic_score_codex":0.005790713,"about_ca_topic_score_gemma":0.00781743,"teacher_disagreement_score":0.005790713,"about_ca_system_score_codex":0.0011331933,"about_ca_system_score_gemma":0.0023017337,"threshold_uncertainty_score":0.017775357},"labels":[],"label_agreement":null},{"id":"W2579722971","doi":"10.7287/peerj.preprints.1705v1","title":"The unreasonable effectiveness of traditional information retrieval in crash report deduplication","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Crash; Computer science; Data deduplication; Software; Information retrieval; Scalability; Precision and recall; Set (abstract data type); Database; Data science; World Wide Web; Software engineering; Operating system","score_opus":0.015265970731580713,"score_gpt":0.2492900997646688,"score_spread":0.23402412903308809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2579722971","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53172815,0.04095521,0.35745066,0.006870482,0.0020930173,0.0015390507,0.00619617,0.0314094,0.02175775],"genre_scores_gemma":[0.6811616,0.006312379,0.29667437,0.0013489512,0.00058661326,0.00029327418,0.005778881,0.0009282769,0.0069156364],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9896402,0.002831167,0.0012938315,0.0017520989,0.003981525,0.0005011604],"domain_scores_gemma":[0.9397359,0.033910356,0.0027656755,0.016785968,0.0062937094,0.00050835265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01226868,0.0014046076,0.0019198285,0.0059565324,0.0018534528,0.0045128483,0.0031991636,0.0021277403,0.0017409248],"category_scores_gemma":[0.055410847,0.00069924956,0.00092993944,0.0068978905,0.0017088328,0.010702564,0.002119615,0.001726293,0.0037017548],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013888711,0.00058770826,0.014645534,0.002703605,0.00044814264,0.0004463361,0.0010649869,0.022437092,0.040026277,0.0048330175,0.05123869,0.86017966],"study_design_scores_gemma":[0.0007431934,0.003742386,0.046286874,0.0010599336,0.001165324,0.008314438,0.0054498855,0.37656668,0.37691513,0.037632488,0.14132991,0.0007937249],"about_ca_topic_score_codex":0.004745031,"about_ca_topic_score_gemma":0.0062389793,"teacher_disagreement_score":0.01226868,"about_ca_system_score_codex":0.0013932934,"about_ca_system_score_gemma":0.0022068669,"threshold_uncertainty_score":0.06488377},"labels":[],"label_agreement":null},{"id":"W2579754794","doi":"10.1109/icsme.2016.83","title":"Why are Commits Being Reverted?: A Comparative Study of Industrial and Open Source Projects","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Commit; Codebase; Computer science; Process (computing); Upload; Process management; Business; Software; Risk analysis (engineering); Computer security; Operations management; World Wide Web; Engineering; Database; Operating system","score_opus":0.12434100384122214,"score_gpt":0.3316756394238308,"score_spread":0.20733463558260867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2579754794","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99844617,0.0001659146,0.00025941816,0.00017663135,0.0000043545656,0.000017816707,0.000044071297,0.000005646422,0.0008799234],"genre_scores_gemma":[0.9993236,0.00015435003,0.00019210289,0.000035006193,0.000008366458,0.000020499221,0.00006202479,0.00000853545,0.00019558228],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9869984,0.0068347487,0.0010276942,0.0011955738,0.0029949367,0.00094872486],"domain_scores_gemma":[0.86020154,0.08952855,0.030941362,0.002766231,0.011817093,0.0047452827],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012241569,0.00023398406,0.00036257712,0.005113011,0.0017276198,0.0031091503,0.0011236897,0.000991923,0.0016970428],"category_scores_gemma":[0.10104459,0.00033009765,0.0003267892,0.005012826,0.002390324,0.0046812072,0.002721313,0.0011788742,0.00029398967],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038558483,0.00040436813,0.66335547,0.00043776486,0.00010877008,0.001237173,0.28822678,0.00034866028,0.0014744287,0.0016534345,0.0008995691,0.04146802],"study_design_scores_gemma":[0.000014923057,0.00044526428,0.7199634,0.00016114519,0.000025888761,0.00052068906,0.27396604,0.0007639198,0.00026410577,0.0004873571,0.0033447621,0.000042525375],"about_ca_topic_score_codex":0.005020421,"about_ca_topic_score_gemma":0.008665301,"teacher_disagreement_score":0.012241569,"about_ca_system_score_codex":0.0020248515,"about_ca_system_score_gemma":0.0015690185,"threshold_uncertainty_score":0.06474042},"labels":[],"label_agreement":null},{"id":"W2584896789","doi":"10.1016/j.scico.2016.11.006","title":"Projecting programs on specifications: Definition and implications","year":2017,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Computer science; Programming language; Operator (biology); Specification language; Projection (relational algebra); Theoretical computer science; Algorithm","score_opus":0.11570300031759101,"score_gpt":0.329896731864103,"score_spread":0.214193731546512,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2584896789","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034861848,0.00037457116,0.9464621,0.0032873545,0.000071319075,0.00009267087,0.0002459951,0.00042705217,0.01417717],"genre_scores_gemma":[0.4840581,0.0017976882,0.506831,0.0007365178,0.00019418441,0.0004974561,0.00079145026,0.0003728422,0.0047208667],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9918194,0.004140243,0.000607403,0.00096447725,0.0020242229,0.00044439864],"domain_scores_gemma":[0.9626166,0.023379676,0.0018482588,0.004758313,0.006066201,0.0013309676],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007134377,0.0010451818,0.0009947441,0.0019550917,0.0016760118,0.004846658,0.0030097927,0.002640545,0.0067914757],"category_scores_gemma":[0.044089034,0.0010593816,0.0016536044,0.004165896,0.009196509,0.016617978,0.0075081964,0.005818884,0.0008228128],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008021081,0.000012846556,0.0005313744,0.0000411692,0.000005544771,0.000022383521,0.00023385436,0.002237637,0.00015345542,0.98804456,0.00030131626,0.008407893],"study_design_scores_gemma":[0.0000093943,0.000020495647,0.0001987305,0.00006049094,0.000008192065,0.00005434565,0.0002149982,0.011949176,0.00045080087,0.9850963,0.0019240346,0.0000130159415],"about_ca_topic_score_codex":0.0028482855,"about_ca_topic_score_gemma":0.0017422248,"teacher_disagreement_score":0.007134377,"about_ca_system_score_codex":0.0014714439,"about_ca_system_score_gemma":0.003092776,"threshold_uncertainty_score":0.037730634},"labels":[],"label_agreement":null},{"id":"W2587703861","doi":"10.1002/smr.1842","title":"The relationship between evolutionary coupling and defects in large industrial software","year":2017,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Engineering and Physical Sciences Research Council; Natural Sciences and Engineering Research Council of Canada; Türkiye Bilimler Akademisi; Türkiye Bilimsel ve Teknolojik Araştırma Kurumu","keywords":"Software; Software evolution; Software system; Computer science; Correlation; Coupling (piping); Process (computing); Data mining; Software engineering; Software construction; Engineering; Mathematics; Programming language","score_opus":0.05343133578773611,"score_gpt":0.3163491672404003,"score_spread":0.2629178314526642,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2587703861","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9984754,0.00007446756,0.001154934,0.000020761672,0.0000011191789,0.000005617202,0.00006868384,0.000013517403,0.00018551589],"genre_scores_gemma":[0.9994437,0.000017157185,0.00039154955,0.000002644775,0.0000019063374,0.0000042432066,0.00009287522,0.000004092602,0.00004182703],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9944811,0.002011018,0.00065488246,0.0007636852,0.0017935183,0.00029582428],"domain_scores_gemma":[0.78648084,0.13682698,0.053578675,0.0076133553,0.013385372,0.0021147793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058355494,0.00036634834,0.00028253923,0.005477461,0.00029709755,0.0009260062,0.00066568394,0.00047750416,0.00078138354],"category_scores_gemma":[0.06564411,0.0003110906,0.0003971374,0.0040413393,0.00075521274,0.0013848335,0.0012382511,0.0007446762,0.000121374986],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040893916,0.00004895626,0.9910908,0.000020026871,0.000058758433,0.00010627862,0.0002518062,0.0024379583,0.0003605258,0.00014596694,0.000046427784,0.00539162],"study_design_scores_gemma":[0.0000034069524,0.00008478012,0.9892161,0.000014239811,0.00002413709,0.00017868621,0.00019390213,0.009413889,0.0004346299,0.0002983501,0.0001287624,0.000009198297],"about_ca_topic_score_codex":0.003250104,"about_ca_topic_score_gemma":0.0028333946,"teacher_disagreement_score":0.0058355494,"about_ca_system_score_codex":0.0007124419,"about_ca_system_score_gemma":0.00046796727,"threshold_uncertainty_score":0.030861676},"labels":[],"label_agreement":null},{"id":"W2588556346","doi":"10.1142/s0218194016400106","title":"A Machine Learning Based Approach for Evaluating Clone Detection Tools for a Generalized and Accurate Precision","year":2016,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Generality; clone (Java method); Computer science; Measure (data warehouse); Machine learning; Software; Variety (cybernetics); Data mining; Artificial intelligence; Java; Sample (material); Detector; Algorithm; Programming language","score_opus":0.03570845271944944,"score_gpt":0.30564613767927035,"score_spread":0.2699376849598209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2588556346","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07149266,0.0005773662,0.9190514,0.00023718686,0.00007359544,0.0004838298,0.00058134063,0.0039614267,0.0035410912],"genre_scores_gemma":[0.43931803,0.0001353912,0.55745226,0.00016963927,0.00006594992,0.0005635606,0.00076859904,0.00022137571,0.0013051712],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9740776,0.004728216,0.0028785986,0.0049299737,0.012570876,0.0008148107],"domain_scores_gemma":[0.92910343,0.03541596,0.010086668,0.009851192,0.014857093,0.00068574556],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01296687,0.0020899174,0.0022735146,0.011697088,0.0013383267,0.0038760684,0.0029208728,0.0031836214,0.0014477748],"category_scores_gemma":[0.07361379,0.0006050028,0.0017996724,0.0075449557,0.0016132264,0.004794132,0.0021300258,0.0024266173,0.0009980486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007225028,0.0006419948,0.071093075,0.0007397162,0.00085052,0.00023785386,0.0006499783,0.16278428,0.031718675,0.012797764,0.0038547434,0.7139089],"study_design_scores_gemma":[0.00006629855,0.0006987477,0.022184595,0.00012446412,0.00018873105,0.00047956032,0.00017647489,0.9261977,0.028948188,0.017106717,0.0036731258,0.00015544296],"about_ca_topic_score_codex":0.0030097151,"about_ca_topic_score_gemma":0.0037189813,"teacher_disagreement_score":0.9870331,"about_ca_system_score_codex":0.002496856,"about_ca_system_score_gemma":0.0018880366,"threshold_uncertainty_score":0.06857622},"labels":[],"label_agreement":null},{"id":"W2588976403","doi":"10.7287/peerj.preprints.2723v1","title":"Lifting the curse of stringly-typed code","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"JavaScript; Computer science; String (physics); Programming language; Parsing; Code (set theory); Source code; Rule-based machine translation; Natural language processing; Theoretical computer science; Mathematics; Set (abstract data type)","score_opus":0.04100188388232198,"score_gpt":0.31685430630663497,"score_spread":0.27585242242431296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2588976403","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95825136,0.0008694836,0.028048566,0.0017946806,0.000066843655,0.00009014693,0.0015132093,0.002350415,0.0070153065],"genre_scores_gemma":[0.97038573,0.0004343093,0.02181403,0.00086985336,0.000050249404,0.00005469398,0.0021750536,0.0017419945,0.002474142],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.97973746,0.006463692,0.0014543508,0.0034425026,0.008034259,0.00086774793],"domain_scores_gemma":[0.8246837,0.115065016,0.02074189,0.02534273,0.012697582,0.0014689921],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0117586395,0.00061887916,0.0005825182,0.0042347936,0.0012741982,0.0030278594,0.0012750049,0.0015170461,0.0023971458],"category_scores_gemma":[0.16108419,0.0009442614,0.0006927807,0.004555651,0.0030304645,0.009109887,0.0036386168,0.0028626826,0.0016159196],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093384576,0.0005068876,0.60699123,0.0012833903,0.00041382323,0.0029341958,0.03208171,0.0034434972,0.018862305,0.015013592,0.021812541,0.29572296],"study_design_scores_gemma":[0.000152266,0.0007152109,0.5934934,0.002438797,0.000660461,0.013355401,0.032218236,0.079687215,0.046371643,0.06727152,0.16324598,0.0003898705],"about_ca_topic_score_codex":0.0037580852,"about_ca_topic_score_gemma":0.006724104,"teacher_disagreement_score":0.0117586395,"about_ca_system_score_codex":0.0008260191,"about_ca_system_score_gemma":0.0018930995,"threshold_uncertainty_score":0.06218636},"labels":[],"label_agreement":null},{"id":"W2589254520","doi":"","title":"Link Typing in Hypertext: Defining Conceptual Attributes","year":2013,"lang":"en","type":"article","venue":"Proceedings of the Annual Conference of CAIS / Actes du congrès annuel de l'ACSI","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Link (geometry); Taxonomy (biology); Hypertext; Computer science; Base (topology); Information retrieval; World Wide Web; Mathematics","score_opus":0.02675956477525647,"score_gpt":0.24428001364052865,"score_spread":0.21752044886527216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2589254520","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86568445,0.0003032363,0.12517697,0.00033403537,0.000045407713,0.0002745721,0.00038786198,0.00026699764,0.007526492],"genre_scores_gemma":[0.94958055,0.000105060055,0.048901718,0.00005706424,0.000017239743,0.00015154497,0.0003835852,0.00005056176,0.0007526506],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99411076,0.0031540596,0.000751283,0.0005192838,0.0011710576,0.00029347636],"domain_scores_gemma":[0.93728226,0.042529345,0.0062730215,0.005520543,0.0070202267,0.0013745822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075824964,0.00033232596,0.0003602621,0.0028757285,0.0012367188,0.0045118653,0.0008190957,0.0010879645,0.0018943255],"category_scores_gemma":[0.052797515,0.00031724825,0.0004073791,0.0034230992,0.0017792926,0.007869563,0.0021422573,0.0012849002,0.00043303563],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001244366,0.0006987317,0.39867708,0.0011818839,0.00009639582,0.0007972677,0.07339146,0.010018027,0.023028761,0.17272915,0.0023813574,0.31575558],"study_design_scores_gemma":[0.0002683383,0.002393566,0.33293188,0.0018295689,0.0004595596,0.0046049366,0.10747868,0.14183442,0.030046793,0.2662796,0.11145924,0.0004134335],"about_ca_topic_score_codex":0.0013076862,"about_ca_topic_score_gemma":0.0011891333,"teacher_disagreement_score":0.0075824964,"about_ca_system_score_codex":0.0010434939,"about_ca_system_score_gemma":0.001014426,"threshold_uncertainty_score":0.040100574},"labels":[],"label_agreement":null},{"id":"W2591512809","doi":"10.1016/j.jss.2017.02.021","title":"Predicting bug-fixing time: A replication study using an open source software project","year":2017,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Replication (statistics); Replicate; Computer science; Software bug; Software; Open source; Statistics; Operating system; Mathematics","score_opus":0.07000649680391977,"score_gpt":0.350390490200443,"score_spread":0.2803839933965232,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2591512809","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9976641,0.000064728876,0.0015367238,0.000045024535,0.000022583517,0.00013511839,0.00021812752,0.000088048895,0.00022566093],"genre_scores_gemma":[0.9895649,0.00008593881,0.008131081,0.00006848072,0.000026380478,0.00029796697,0.0012189534,0.00008532029,0.0005209701],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.988305,0.0053279693,0.0012905957,0.0022406594,0.0024808524,0.00035498783],"domain_scores_gemma":[0.81469476,0.0966474,0.01419275,0.03314009,0.038745098,0.0025799274],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.017100107,0.001286303,0.0010540778,0.0028781483,0.0023945598,0.0014857755,0.0029333043,0.0024277824,0.00087149686],"category_scores_gemma":[0.10330294,0.0009353933,0.0021583016,0.0030460272,0.0011072467,0.002492814,0.0013964036,0.0020465096,0.00058979256],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024208927,0.0096075125,0.8830478,0.000601309,0.0013297969,0.0012689375,0.009995297,0.008272018,0.011543729,0.00031572828,0.0020781786,0.06951884],"study_design_scores_gemma":[0.0012046797,0.018418688,0.8862948,0.00023651883,0.0024132999,0.0014634298,0.009220652,0.06144613,0.013071257,0.00095654017,0.0048690154,0.00040507386],"about_ca_topic_score_codex":0.021300355,"about_ca_topic_score_gemma":0.024479732,"teacher_disagreement_score":0.9828999,"about_ca_system_score_codex":0.0016181283,"about_ca_system_score_gemma":0.0026440308,"threshold_uncertainty_score":0.09043509},"labels":[],"label_agreement":null},{"id":"W2591557490","doi":"","title":"Predicting Method Crashes with Bytecode Operations","year":2013,"lang":"en","type":"article","venue":"ACM International Conference Proceeding Series","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"AspectJ; Crash; Computer science; Bytecode; Overhead (engineering); Software; Eclipse; Real-time computing; Machine learning; Data mining; Operating system; Aspect-oriented programming; Java","score_opus":0.038946884731466944,"score_gpt":0.309919357174983,"score_spread":0.270972472443516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2591557490","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9697077,0.0002777513,0.024017336,0.00015974454,0.00004574733,0.00006432834,0.0012841118,0.0037500628,0.0006931149],"genre_scores_gemma":[0.98125607,0.000118935386,0.015675774,0.000023256624,0.00002175935,0.000036682908,0.0022370392,0.00010618409,0.0005243258],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99798155,0.00024688878,0.00020689503,0.00052812975,0.00084114674,0.00019545788],"domain_scores_gemma":[0.96718365,0.016784199,0.006588169,0.0029821491,0.005260351,0.0012015254],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022892023,0.0014121156,0.00058504596,0.0035759453,0.0004098543,0.000953818,0.0009900635,0.0010248814,0.0006357327],"category_scores_gemma":[0.022972435,0.000441401,0.0006967557,0.001756136,0.000383119,0.0015330521,0.00080045307,0.0013206098,0.0006475933],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003904008,0.00029442675,0.84861094,0.0001093038,0.00012476495,0.00034117722,0.000274167,0.046902362,0.0053680777,0.00023047048,0.002960107,0.09439386],"study_design_scores_gemma":[0.000016341444,0.00026646568,0.15303203,0.000025919082,0.000079839665,0.00042517186,0.00017905227,0.8329962,0.010929053,0.00096317707,0.0010554759,0.00003124283],"about_ca_topic_score_codex":0.006994943,"about_ca_topic_score_gemma":0.009378837,"teacher_disagreement_score":0.006994943,"about_ca_system_score_codex":0.0005500855,"about_ca_system_score_gemma":0.00078085996,"threshold_uncertainty_score":0.013908446},"labels":[],"label_agreement":null},{"id":"W2591607085","doi":"","title":"An effective method for detecting duplicate crash reports using crash traces and hidden Markov models","year":2016,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Crash; Computer science; Hidden Markov model; Software; Task (project management); Eclipse; Data mining; Machine learning; Artificial intelligence; Operating system; Engineering","score_opus":0.017807790907224372,"score_gpt":0.28377816607998235,"score_spread":0.26597037517275796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2591607085","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02079778,0.00041321528,0.97448266,0.00013536731,0.000070668764,0.0001439652,0.00026296204,0.0033316128,0.00036172487],"genre_scores_gemma":[0.37382638,0.0004635926,0.6213073,0.00017190003,0.00013825919,0.00026241288,0.0015349516,0.00028226693,0.0020128884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.996879,0.0005863001,0.00032114433,0.0006292205,0.0013532775,0.00023107941],"domain_scores_gemma":[0.98971736,0.004308424,0.001288878,0.0016791874,0.0026844293,0.0003217221],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025248649,0.0014823723,0.0016236432,0.004678395,0.0008315928,0.0011878969,0.0028305585,0.0017734363,0.00084734533],"category_scores_gemma":[0.011268124,0.0009355716,0.00125187,0.0021384992,0.0005751005,0.0023840521,0.0018909448,0.0019426265,0.00082247413],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053324783,0.0006575204,0.043861184,0.00050682056,0.00046592328,0.0014489383,0.00074641523,0.060842674,0.035084173,0.00597674,0.00901154,0.84086484],"study_design_scores_gemma":[0.00004893354,0.00013643144,0.0058893855,0.000043344033,0.00020338834,0.0012567604,0.0001641522,0.95755905,0.024330223,0.0061148033,0.0041343416,0.00011916183],"about_ca_topic_score_codex":0.0063608005,"about_ca_topic_score_gemma":0.0067961444,"teacher_disagreement_score":0.0063608005,"about_ca_system_score_codex":0.00064539455,"about_ca_system_score_gemma":0.0021707027,"threshold_uncertainty_score":0.013352931},"labels":[],"label_agreement":null},{"id":"W2591663134","doi":"10.3233/idt-170284","title":"Planning for the next software release using adaptive network-based fuzzy inference system","year":2017,"lang":"en","type":"article","venue":"Intelligent Decision Technologies","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Adaptive neuro fuzzy inference system; Inference; Process (computing); Computer science; Software release life cycle; Fuzzy logic; Data mining; Software; Inference engine; Machine learning; Reliability (semiconductor); Fuzzy inference system; Perspective (graphical); Artificial intelligence; Fuzzy control system; Software quality; Software development","score_opus":0.14157031471738002,"score_gpt":0.3592322270332248,"score_spread":0.2176619123158448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2591663134","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047942735,0.00030904703,0.9466157,0.00019922771,0.000041474083,0.00013661142,0.00008746201,0.0006492927,0.004018472],"genre_scores_gemma":[0.8696508,0.00029577158,0.12739357,0.00006694182,0.00002555537,0.0002707018,0.00018963024,0.000026090338,0.0020808869],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994746,0.00011242248,0.000053781303,0.00015690431,0.00014085067,0.000061585226],"domain_scores_gemma":[0.9991954,0.000490964,0.00011691872,0.00001910354,0.00015327941,0.000024281022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010055447,0.00075033185,0.0007242903,0.0007217818,0.00064365,0.0010414416,0.0009639384,0.0009345916,0.001701286],"category_scores_gemma":[0.0019773515,0.00036023837,0.00065665285,0.00048440055,0.00036764666,0.0007200204,0.0004973021,0.00079106103,0.00019644563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010529173,0.00005412799,0.000901462,0.00008342338,0.000037067544,0.00012936328,0.00011211977,0.94936746,0.0023224212,0.0013762084,0.00040700487,0.04510415],"study_design_scores_gemma":[0.000006189734,0.000017748835,0.000121883095,0.000006220915,0.000009635675,0.0000074874224,0.000009092022,0.99890316,0.00036246196,0.0004291302,0.00012267233,0.00000423731],"about_ca_topic_score_codex":0.01926824,"about_ca_topic_score_gemma":0.01833879,"teacher_disagreement_score":0.01926824,"about_ca_system_score_codex":0.0011025234,"about_ca_system_score_gemma":0.0012669216,"threshold_uncertainty_score":0.038312197},"labels":[],"label_agreement":null},{"id":"W2592568457","doi":"10.1002/smr.1843","title":"MORE: A multi‐objective refactoring recommendation approach to introducing design patterns and fixing code smells","year":2017,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Japan Society for the Promotion of Science; Science Foundation Ireland","keywords":"Code refactoring; Code smell; Computer science; Software quality; Software engineering; Maintainability; Class (philosophy); Software maintenance; Source code; Quality (philosophy); Software; Software system; Software development; Programming language; Artificial intelligence","score_opus":0.044680821213054306,"score_gpt":0.3188694893080094,"score_spread":0.2741886680949551,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2592568457","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045900468,0.00054744503,0.9471709,0.000363327,0.00005580995,0.00034939317,0.00019829944,0.0026171545,0.0027972786],"genre_scores_gemma":[0.28987023,0.00031075053,0.70290446,0.00038362417,0.00004819808,0.0005283956,0.00062823226,0.00031588177,0.005010202],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982462,0.00058453676,0.00011440473,0.00031645203,0.00061151385,0.00012688084],"domain_scores_gemma":[0.9968618,0.0015821257,0.00042124826,0.00029244844,0.00071808335,0.00012425585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023494107,0.001819756,0.0013225727,0.0027406344,0.000553049,0.00091065286,0.0021821973,0.0015380203,0.002752136],"category_scores_gemma":[0.005115537,0.0007672471,0.0016671591,0.001453804,0.000521513,0.0010886991,0.0009248582,0.00096571544,0.0005678128],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001947964,0.0006784819,0.0059831413,0.00053041725,0.00036408033,0.00023556176,0.00039576177,0.49231827,0.012966973,0.004199106,0.0046906015,0.47744283],"study_design_scores_gemma":[0.000067246445,0.00024112902,0.0012394628,0.000037809572,0.00011753864,0.00008178122,0.00006975247,0.9908647,0.0022924345,0.0024444463,0.0025138229,0.00002992895],"about_ca_topic_score_codex":0.008561646,"about_ca_topic_score_gemma":0.01617258,"teacher_disagreement_score":0.008561646,"about_ca_system_score_codex":0.0010971361,"about_ca_system_score_gemma":0.0019746819,"threshold_uncertainty_score":0.017023623},"labels":[],"label_agreement":null},{"id":"W2593117527","doi":"10.1007/s11219-017-9361-y","title":"An empirical study of crash-inducing commits in Mozilla Firefox","year":2017,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Crash; Computer science; Empirical research; Software engineering; Business; Programming language; Mathematics; Statistics","score_opus":0.09836511809479921,"score_gpt":0.4380468293221385,"score_spread":0.3396817112273393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2593117527","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9991968,0.00002317268,0.000109584034,0.000058257727,0.000002845241,0.000013620101,0.000055761837,0.000008122755,0.00053183537],"genre_scores_gemma":[0.9991462,0.00002695476,0.00018611507,0.000025946181,0.000004530247,0.00001338903,0.00015950039,0.0000062662934,0.00043100462],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9978446,0.00077002397,0.00018073866,0.000239067,0.0007169374,0.00024866007],"domain_scores_gemma":[0.8925157,0.05750228,0.028372575,0.0038924539,0.013139288,0.004577668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034020368,0.00023709146,0.00020468874,0.0018172882,0.001188373,0.0010274075,0.0008955069,0.0009777728,0.0029922812],"category_scores_gemma":[0.056456927,0.0003129886,0.00018602186,0.0012876488,0.0006934627,0.0012341834,0.0006467877,0.0018575585,0.00065236964],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003258003,0.0016009773,0.9776649,0.000054087628,0.000041061845,0.00038215367,0.005382492,0.00037446886,0.00092131674,0.00047085522,0.0010443819,0.011737539],"study_design_scores_gemma":[0.000028243134,0.00063745416,0.9861172,0.000038822312,0.00002730905,0.00022884959,0.008721125,0.0025291138,0.0004617306,0.00018926394,0.0010030622,0.000017889945],"about_ca_topic_score_codex":0.017659115,"about_ca_topic_score_gemma":0.031974994,"teacher_disagreement_score":0.017659115,"about_ca_system_score_codex":0.0012867214,"about_ca_system_score_gemma":0.0013738216,"threshold_uncertainty_score":0.03511268},"labels":[],"label_agreement":null},{"id":"W2593891526","doi":"10.1109/icse-nier.2017.12","title":"Building Usage Profiles Using Deep Neural Nets","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Profiling (computer programming); Software; Convolutional neural network; Task (project management); Artificial neural network; Construct (python library); Deep learning","score_opus":0.0497947721742919,"score_gpt":0.3320615793785307,"score_spread":0.2822668072042388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2593891526","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5089695,0.0010289837,0.459141,0.0004996992,0.00011108392,0.0002829415,0.009612523,0.011019788,0.009334451],"genre_scores_gemma":[0.878897,0.00039348553,0.10230497,0.00015044736,0.000034458226,0.00017148544,0.01352223,0.00027110684,0.0042548124],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993631,0.00009993724,0.0000557675,0.00022116842,0.00014932727,0.000110790745],"domain_scores_gemma":[0.99842775,0.0005341184,0.00025473625,0.00017979257,0.0004843523,0.000119336466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005357568,0.0011683472,0.00055900024,0.0033664813,0.00032238424,0.0007789518,0.0006897003,0.00068539794,0.0015081764],"category_scores_gemma":[0.0040754993,0.00051847094,0.00054717826,0.0017770699,0.00016252705,0.001728152,0.00068314583,0.0009958223,0.0015155734],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059091055,0.00081807014,0.10736577,0.0002720429,0.00021429535,0.00044994013,0.00044110286,0.1434527,0.011892685,0.0023552997,0.015974976,0.7161723],"study_design_scores_gemma":[0.0000066189596,0.000048298258,0.0099003585,0.00002905695,0.00001709309,0.000057039982,0.00007916862,0.98248094,0.0033045853,0.002372643,0.0016888838,0.000015372761],"about_ca_topic_score_codex":0.02005365,"about_ca_topic_score_gemma":0.037873648,"teacher_disagreement_score":0.02005365,"about_ca_system_score_codex":0.00109874,"about_ca_system_score_gemma":0.00072800013,"threshold_uncertainty_score":0.03987384},"labels":[],"label_agreement":null},{"id":"W2594091530","doi":"","title":"Automatic prediction of the severity of bugs using stack traces","year":2016,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software bug; Stack (abstract data type); Crash; Eclipse; Data mining; Function (biology); Process (computing); Software regression; Machine learning; Programming language; Software; Software quality; Software development","score_opus":0.015591502299441106,"score_gpt":0.23841786484292785,"score_spread":0.22282636254348673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2594091530","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8923774,0.0011153391,0.06969551,0.00031713006,0.00010545683,0.00026462134,0.013163061,0.021469742,0.0014917152],"genre_scores_gemma":[0.9192828,0.00046296103,0.05375337,0.000038925915,0.0000670617,0.000106160114,0.024740271,0.00039284752,0.0011557315],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983027,0.00023317101,0.00023514216,0.0003655379,0.00072430336,0.00013911437],"domain_scores_gemma":[0.98478323,0.0051442184,0.0041419454,0.0012312318,0.0039032532,0.00079616753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013821411,0.0013831945,0.0007710081,0.013848252,0.0002961245,0.00096571486,0.0007661789,0.0008265576,0.0007341789],"category_scores_gemma":[0.012535177,0.0005475858,0.00076774036,0.0032957038,0.00022370712,0.0013619874,0.00072560913,0.00069458614,0.0008707822],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006350731,0.00049035536,0.5227391,0.0006257378,0.0003162595,0.0012942955,0.00081744255,0.02950151,0.045547087,0.00073118566,0.012176769,0.3851252],"study_design_scores_gemma":[0.000081156184,0.0005692823,0.31434932,0.00012873001,0.0002367423,0.0008641671,0.00039596693,0.64072615,0.03341062,0.0018535435,0.007246425,0.00013790066],"about_ca_topic_score_codex":0.013655308,"about_ca_topic_score_gemma":0.018290058,"teacher_disagreement_score":0.013848252,"about_ca_system_score_codex":0.00056160585,"about_ca_system_score_gemma":0.0010155757,"threshold_uncertainty_score":0.027151644},"labels":[],"label_agreement":null},{"id":"W2594215239","doi":"","title":"Local versus global models for effort-aware defect prediction","year":2016,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Context (archaeology); Set (abstract data type); Machine learning; Training set; Data mining; Predictive modelling; Artificial intelligence; Data set; Software","score_opus":0.01854812757945171,"score_gpt":0.24519378507215867,"score_spread":0.22664565749270696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2594215239","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6444716,0.0049822438,0.3342428,0.0015697314,0.00024191558,0.0003054765,0.0048781894,0.0042893733,0.005018625],"genre_scores_gemma":[0.953137,0.00057085877,0.038179234,0.00019624899,0.00012435195,0.00018778231,0.004614775,0.00016621558,0.0028236215],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99814963,0.0004817664,0.00012733038,0.0007604689,0.00029088123,0.00018986217],"domain_scores_gemma":[0.9945117,0.0028687913,0.00073497486,0.0007269395,0.00088442094,0.00027318348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004318264,0.002395469,0.0017625843,0.002883929,0.00032502066,0.0014036788,0.0018902611,0.001308554,0.0012745241],"category_scores_gemma":[0.007510382,0.0004162982,0.0017177443,0.001726616,0.0005274676,0.002524669,0.0016258197,0.0019075822,0.00090630446],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062886055,0.0006322877,0.13530809,0.00036910415,0.0006527619,0.0003146908,0.00042665337,0.67217684,0.0029477954,0.0022213438,0.0077681686,0.17655343],"study_design_scores_gemma":[0.000018609991,0.00020287631,0.015140773,0.000048147886,0.00009275952,0.000080968275,0.00011141742,0.9800388,0.000675306,0.0026056133,0.0009566137,0.00002817815],"about_ca_topic_score_codex":0.0076596923,"about_ca_topic_score_gemma":0.011061823,"teacher_disagreement_score":0.0076596923,"about_ca_system_score_codex":0.00063729263,"about_ca_system_score_gemma":0.0007247025,"threshold_uncertainty_score":0.02283746},"labels":[],"label_agreement":null},{"id":"W2594297796","doi":"","title":"An exploratory study on change suggestions for methods using clone detection","year":2016,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Ranking (information retrieval); Computer science; Change detection; Precision and recall; Rank (graph theory); Change analysis; Information retrieval; Data mining; Complement (music); Software evolution; Recall; Data science; Software; Machine learning; Artificial intelligence; Software system; Cognitive psychology; Mathematics","score_opus":0.08106462202636616,"score_gpt":0.3533319726558519,"score_spread":0.27226735062948576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2594297796","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98669213,0.00053910224,0.009952552,0.00019995391,0.000023768709,0.00050088565,0.0003920082,0.00033018555,0.0013693414],"genre_scores_gemma":[0.96095395,0.00040232888,0.036206666,0.00019394979,0.00003938661,0.00047994952,0.00063124136,0.00013354837,0.0009589607],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9742717,0.01529847,0.002271695,0.002567544,0.005034269,0.0005562201],"domain_scores_gemma":[0.49020067,0.44405156,0.023505215,0.014480511,0.02635918,0.0014028559],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022518555,0.0007426823,0.00076629606,0.0054477425,0.0011506913,0.0021808033,0.0015671373,0.001273816,0.0010435907],"category_scores_gemma":[0.21994765,0.0006806056,0.0006430727,0.0034272142,0.00091704604,0.003676091,0.0014093885,0.0016134529,0.00036538357],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001783549,0.0026025977,0.5035568,0.0046164626,0.00038794978,0.0020053554,0.13905832,0.0017118527,0.025269276,0.001188807,0.0029659865,0.3148529],"study_design_scores_gemma":[0.0004104981,0.011819808,0.7660219,0.0021121704,0.000971736,0.005447177,0.07848397,0.04424025,0.047324523,0.0020924304,0.040580988,0.00049455627],"about_ca_topic_score_codex":0.0025917138,"about_ca_topic_score_gemma":0.0050436575,"teacher_disagreement_score":0.022518555,"about_ca_system_score_codex":0.0011672499,"about_ca_system_score_gemma":0.0014492383,"threshold_uncertainty_score":0.119090974},"labels":[],"label_agreement":null},{"id":"W2594779357","doi":"","title":"On the Computational Complexity of Software (Re)Modularization: Elaborations and Opportunities.","year":2015,"lang":"en","type":"article","venue":"BICT","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Modular programming; Computer science; Software; Mathematical proof; Heuristic; Simulated annealing; Theoretical computer science; Software engineering; Programming complexity; Software construction; Software system; Programming language; Algorithm; Mathematics; Artificial intelligence","score_opus":0.17772296183058603,"score_gpt":0.2992789429028447,"score_spread":0.12155598107225865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2594779357","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03832483,0.01421647,0.8648511,0.022133343,0.0005810104,0.00024294907,0.0003793872,0.0004635056,0.058807287],"genre_scores_gemma":[0.5888529,0.028293919,0.3605327,0.0039393757,0.0034058904,0.0008221553,0.00133587,0.000697101,0.0121200895],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9951826,0.0017180499,0.0002186331,0.0009308914,0.0013852937,0.0005645609],"domain_scores_gemma":[0.9302804,0.060217004,0.0020852974,0.0052592126,0.0014371813,0.0007209341],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007168542,0.0015486408,0.0013733389,0.002238411,0.0018178973,0.005206112,0.003955742,0.0030640853,0.01004495],"category_scores_gemma":[0.052157745,0.0012379191,0.0025642335,0.0034981456,0.010391359,0.01734451,0.0055338186,0.010163645,0.0013652926],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000674972,0.00008596292,0.0008093607,0.0005591381,0.00006023099,0.00015417808,0.00023980282,0.06986713,0.0008532939,0.87231433,0.009060993,0.045928054],"study_design_scores_gemma":[0.00001866079,0.0000307101,0.00035450666,0.000106766405,0.000024217352,0.00015188857,0.000082661616,0.09456615,0.00053405203,0.89741534,0.0066947546,0.000020177664],"about_ca_topic_score_codex":0.0019711384,"about_ca_topic_score_gemma":0.0022483927,"teacher_disagreement_score":0.01004495,"about_ca_system_score_codex":0.003940564,"about_ca_system_score_gemma":0.0027024767,"threshold_uncertainty_score":0.037911355},"labels":[],"label_agreement":null},{"id":"W2597809331","doi":"10.1109/msr.2017.65","title":"Extracting Build Changes with BUILDDIFF","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Code refactoring; Computer science; Java; Software evolution; Software; Precision and recall; Software maintenance; Software engineering; Software system; Data mining; Machine learning; Software construction; Programming language","score_opus":0.029416569007886112,"score_gpt":0.29435521294549977,"score_spread":0.2649386439376137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2597809331","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2364218,0.0032308681,0.5542723,0.0006505476,0.0005918535,0.0012696979,0.04963889,0.14506143,0.0088626575],"genre_scores_gemma":[0.285938,0.0010392328,0.5944442,0.00028230157,0.00018215855,0.0012666202,0.10231038,0.009578642,0.0049584466],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9937266,0.00068339094,0.0010376852,0.0017017947,0.0026202672,0.00023014906],"domain_scores_gemma":[0.968888,0.011396485,0.0052404185,0.007988376,0.006056265,0.00043055252],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003033686,0.001801364,0.00092510064,0.013276486,0.0007627233,0.0020983815,0.0015508123,0.0013512174,0.0017670097],"category_scores_gemma":[0.028875906,0.0012400179,0.0013324413,0.0065793577,0.0005459483,0.0036858455,0.0029758418,0.0017488793,0.0019771878],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043366934,0.0002645401,0.14686838,0.0022706946,0.00040508484,0.00140114,0.0035267514,0.009527278,0.030224498,0.0037592142,0.048490718,0.752828],"study_design_scores_gemma":[0.00018093098,0.0007150897,0.27370605,0.0006718909,0.00074310217,0.0041622026,0.0019561113,0.22674167,0.13206707,0.014761274,0.34374827,0.0005462554],"about_ca_topic_score_codex":0.0031269274,"about_ca_topic_score_gemma":0.0053772638,"teacher_disagreement_score":0.013276486,"about_ca_system_score_codex":0.00073910935,"about_ca_system_score_gemma":0.0013444192,"threshold_uncertainty_score":0.016043842},"labels":[],"label_agreement":null},{"id":"W2598761292","doi":"10.1109/iwsc.2017.7880507","title":"Does cloned code increase maintenance effort?","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Code (set theory); Programming language","score_opus":0.014813332863914744,"score_gpt":0.27691906136334016,"score_spread":0.2621057284994254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2598761292","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9824316,0.0013606638,0.012761486,0.0005494717,0.00003936209,0.0000594665,0.00015686637,0.00059398287,0.0020471576],"genre_scores_gemma":[0.989507,0.00043502793,0.008745118,0.00010731493,0.00003099445,0.000026952983,0.00023568535,0.00015506939,0.0007568641],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99260956,0.002388974,0.000628485,0.0010459124,0.0029575091,0.00036964755],"domain_scores_gemma":[0.829917,0.11627323,0.032257862,0.009749758,0.010230061,0.0015720169],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037937148,0.0007057072,0.0003863528,0.0024265756,0.00035691482,0.001699493,0.0008811455,0.0009826439,0.0014099516],"category_scores_gemma":[0.07805141,0.0005221805,0.0006457557,0.0014680723,0.0007404594,0.002749236,0.0009104294,0.0007357306,0.00037353663],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000584162,0.00031533287,0.73096246,0.0008520519,0.00034974175,0.001296135,0.0026749657,0.0061694155,0.030350296,0.0012186656,0.0009570604,0.22426972],"study_design_scores_gemma":[0.00003567947,0.0013013257,0.9373434,0.00027797229,0.00041236298,0.0025379,0.0014084904,0.026530636,0.021557687,0.0031383566,0.005389456,0.00006671806],"about_ca_topic_score_codex":0.0018979098,"about_ca_topic_score_gemma":0.003541351,"teacher_disagreement_score":0.0037937148,"about_ca_system_score_codex":0.00083350705,"about_ca_system_score_gemma":0.0008134817,"threshold_uncertainty_score":0.02006334},"labels":[],"label_agreement":null},{"id":"W2599817676","doi":"10.1145/3034950.3034957","title":"Success or Failure Identification for GitHub's Open Source Projects","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Carleton University","funders":"","keywords":"Java; Computer science; JavaScript; Open source; Identification (biology); Popularity; World Wide Web; Sample (material); Web application; Source code; Software engineering; Information retrieval; Programming language; Data science; Software","score_opus":0.07173796243589714,"score_gpt":0.36350510436341066,"score_spread":0.2917671419275135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2599817676","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96945,0.000876671,0.008124253,0.00043025042,0.000086419925,0.00015605794,0.01620468,0.0017709765,0.002900793],"genre_scores_gemma":[0.92303646,0.00045656782,0.021010274,0.00013765377,0.00005941439,0.0003178566,0.05219468,0.0005653646,0.002221736],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9927135,0.0013626716,0.0008238957,0.0016100479,0.0027172733,0.00077261333],"domain_scores_gemma":[0.96525264,0.017406603,0.00545259,0.004589965,0.0053447383,0.001953469],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053364104,0.0008603877,0.000733158,0.009907325,0.0011573805,0.0018357262,0.0013193595,0.0010927701,0.0014436899],"category_scores_gemma":[0.031454068,0.00024955138,0.000675667,0.007839097,0.00093115587,0.001850727,0.0021377748,0.0012002748,0.0014806258],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005548053,0.00036651583,0.81507224,0.0011871976,0.0002465037,0.0009983496,0.0018842515,0.0058251573,0.0038875237,0.0022835524,0.041750547,0.12594333],"study_design_scores_gemma":[0.0000892777,0.00043593525,0.84669733,0.0003090074,0.00016482483,0.0032866728,0.004534977,0.08028108,0.011353034,0.005217972,0.047444806,0.00018512411],"about_ca_topic_score_codex":0.005541896,"about_ca_topic_score_gemma":0.008431132,"teacher_disagreement_score":0.009907325,"about_ca_system_score_codex":0.0009179155,"about_ca_system_score_gemma":0.0010979541,"threshold_uncertainty_score":0.028221965},"labels":[],"label_agreement":null},{"id":"W2599821978","doi":"10.1109/saner.2017.7884611","title":"STRICT: Information retrieval based search term identification for concept location","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Task (project management); Information retrieval; Term (time); Baseline (sea); Software; Identification (biology); Natural language; Domain (mathematical analysis); Source code; Quality (philosophy); Software maintenance; Software development; Data mining; Natural language processing; Artificial intelligence; Programming language","score_opus":0.04110697748763638,"score_gpt":0.32599615927465847,"score_spread":0.2848891817870221,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2599821978","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13587423,0.007926276,0.78003633,0.0009660205,0.0007022455,0.0022613571,0.015634555,0.046754364,0.009844637],"genre_scores_gemma":[0.19906023,0.0013863605,0.771239,0.00029549823,0.00032084915,0.0008854899,0.018411152,0.0008494006,0.0075521646],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976694,0.00053358526,0.00032408655,0.00038787638,0.00090367603,0.0001813585],"domain_scores_gemma":[0.9941356,0.0028132766,0.00076928193,0.00075123017,0.0012913962,0.00023922941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001826201,0.0019411471,0.0016690904,0.01638972,0.0010061748,0.0015307189,0.0016392985,0.0013777333,0.006158333],"category_scores_gemma":[0.010217547,0.00047680349,0.0013157912,0.008420825,0.00058903784,0.0045285285,0.0020815213,0.0011302684,0.0062938468],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011700842,0.00066907966,0.009829175,0.0029159233,0.00029105003,0.00050528796,0.00088923424,0.0037950212,0.09379467,0.0065235086,0.057286892,0.8223302],"study_design_scores_gemma":[0.0011189489,0.003098234,0.03942129,0.00054173166,0.0010876633,0.006971636,0.0025993946,0.5456613,0.19712973,0.03930707,0.16234496,0.00071802014],"about_ca_topic_score_codex":0.0045018746,"about_ca_topic_score_gemma":0.009280022,"teacher_disagreement_score":0.01638972,"about_ca_system_score_codex":0.0009589,"about_ca_system_score_gemma":0.0023790118,"threshold_uncertainty_score":0.02060163},"labels":[],"label_agreement":null},{"id":"W2600599962","doi":"10.1007/s10270-017-0594-9","title":"Managing design-time uncertainty","year":2017,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Université de Montréal","funders":"","keywords":"Computer science; Uncertainty analysis; ENCODE; Systems engineering; Software engineering; Risk analysis (engineering); Management science; Simulation","score_opus":0.04730855297290635,"score_gpt":0.2795139488709415,"score_spread":0.23220539589803516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2600599962","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10164768,0.0012493372,0.87957335,0.0024362195,0.00014910451,0.00008485945,0.000119881195,0.00056824216,0.014171366],"genre_scores_gemma":[0.92599326,0.000401906,0.07058472,0.00025100482,0.00010551992,0.00006930911,0.000110306166,0.00018519306,0.0022987712],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99331,0.002427586,0.0002452274,0.00077833433,0.0025694785,0.0006693696],"domain_scores_gemma":[0.9700827,0.0218501,0.0022461172,0.0030152407,0.0019166254,0.00088913465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067145056,0.0009423879,0.001043719,0.0012251087,0.00091924926,0.003539191,0.0017246831,0.0016315675,0.003852762],"category_scores_gemma":[0.045857094,0.00078672677,0.00082016806,0.0011921477,0.0013077398,0.0055041662,0.0036425155,0.0029118434,0.00041647602],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034739022,0.00019910377,0.0037544987,0.00023541221,0.00019574461,0.00057575625,0.0007195697,0.70193404,0.006543455,0.11049648,0.0029563177,0.17204228],"study_design_scores_gemma":[0.000023492334,0.000084720734,0.0004928005,0.000032165994,0.00006656983,0.00010554329,0.00016179486,0.8500903,0.002341866,0.14204887,0.00453099,0.00002082556],"about_ca_topic_score_codex":0.0011027672,"about_ca_topic_score_gemma":0.0015923892,"teacher_disagreement_score":0.0067145056,"about_ca_system_score_codex":0.002047701,"about_ca_system_score_gemma":0.0024391448,"threshold_uncertainty_score":0.035510123},"labels":[],"label_agreement":null},{"id":"W2600957813","doi":"10.1109/saner.2017.7884630","title":"An empirical study of code smells in JavaScript projects","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"JavaScript; Computer science; Code smell; Programming language; Code (set theory); Empirical research; World Wide Web; Computer security; Software quality; Software; Software development","score_opus":0.08580029150157727,"score_gpt":0.38963270729340443,"score_spread":0.3038324157918272,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2600957813","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99884963,0.00010405794,0.00037475594,0.00008787788,0.0000031490001,0.000023245915,0.00015295176,0.00001557686,0.00038872616],"genre_scores_gemma":[0.9990637,0.00008825419,0.00034430073,0.000027381937,0.0000059986755,0.000028434717,0.00023768964,0.000012319756,0.00019200714],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99344444,0.0017841472,0.00081800483,0.000796684,0.002533336,0.000623428],"domain_scores_gemma":[0.7857242,0.10298631,0.081640154,0.005522895,0.01793958,0.0061867824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067869327,0.0003110992,0.00026957577,0.0026975127,0.0005612655,0.0013350728,0.00067096413,0.0008189166,0.0011553162],"category_scores_gemma":[0.08614527,0.00035271473,0.0004323631,0.0021088338,0.0010260888,0.0022340838,0.0012255552,0.0013005505,0.00040964966],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008898794,0.00013695392,0.98638874,0.00007399647,0.000027982931,0.00019331455,0.004456493,0.00019957584,0.00039984225,0.00006487069,0.0002953514,0.007673812],"study_design_scores_gemma":[0.000004304253,0.00020840229,0.9922999,0.00005503515,0.000013455994,0.00023529422,0.0049076863,0.0012111291,0.00030284753,0.00006853652,0.00067424984,0.000019190247],"about_ca_topic_score_codex":0.0036090054,"about_ca_topic_score_gemma":0.0039949277,"teacher_disagreement_score":0.0067869327,"about_ca_system_score_codex":0.00090254674,"about_ca_system_score_gemma":0.00085332827,"threshold_uncertainty_score":0.035893142},"labels":[],"label_agreement":null},{"id":"W2602038282","doi":"","title":"An Analyses of the McCabe Cyclomatic Complexity Number","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Cyclomatic complexity; Computer science; Software; Software engineering; Software metric; Basis (linear algebra); Software measurement; Data mining; Theoretical computer science; Management science; Software development; Software quality; Mathematics; Programming language; Engineering","score_opus":0.08009049183793308,"score_gpt":0.3707717134687013,"score_spread":0.2906812216307682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2602038282","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2971994,0.0063088476,0.504411,0.008010385,0.0006673656,0.00020317527,0.0011401811,0.00056841504,0.18149118],"genre_scores_gemma":[0.9529136,0.0010995463,0.041542694,0.00026627604,0.00040376358,0.00020536283,0.00019929635,0.0000694802,0.0032999776],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99735284,0.0007484119,0.00013065644,0.00039733722,0.0011281807,0.00024258043],"domain_scores_gemma":[0.98123837,0.011616166,0.0021842662,0.0018420522,0.002452119,0.0006671115],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017538171,0.00047900697,0.00049017847,0.004681347,0.0011977975,0.002408489,0.00054261484,0.00063046697,0.0045374106],"category_scores_gemma":[0.023846537,0.00022839906,0.0006164499,0.0032182643,0.002908813,0.004485535,0.0014803819,0.0014691176,0.00035438873],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026158907,0.000010216901,0.0022774215,0.00006988419,0.000008823384,0.00007216502,0.0002897899,0.0062498106,0.0010645854,0.9776433,0.0019783252,0.010309567],"study_design_scores_gemma":[0.0000069361377,0.00006535587,0.010034411,0.00007512327,0.000015983718,0.0003292289,0.00029251972,0.052011915,0.0023008808,0.916198,0.018603446,0.00006633039],"about_ca_topic_score_codex":0.0019294213,"about_ca_topic_score_gemma":0.0012012922,"teacher_disagreement_score":0.004681347,"about_ca_system_score_codex":0.00232344,"about_ca_system_score_gemma":0.0012572855,"threshold_uncertainty_score":0.016857743},"labels":[],"label_agreement":null},{"id":"W2602793737","doi":"10.1109/msr.2017.50","title":"Rediscovery Datasets: Connecting Duplicate Reports","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software bug; Open source software; Software; Open source; Tracking (education); Tracking system","score_opus":0.030307031169124206,"score_gpt":0.3088519965327439,"score_spread":0.2785449653636197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2602793737","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.102182284,0.003570465,0.018450748,0.0015602802,0.00050243235,0.00075504713,0.85137945,0.01619831,0.005400918],"genre_scores_gemma":[0.051516555,0.00055387965,0.027667984,0.00022403682,0.00008908944,0.0005620255,0.91756827,0.000436029,0.0013821754],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99340546,0.00083722465,0.0013464667,0.0017020333,0.0022294375,0.0004793198],"domain_scores_gemma":[0.9780815,0.0058033303,0.0038831495,0.0061556185,0.0052485596,0.0008278618],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004448,0.0018225687,0.0012109315,0.012892061,0.0016969539,0.0025612935,0.0043123816,0.0030748523,0.0017964124],"category_scores_gemma":[0.027223013,0.0007577231,0.0015305869,0.009278233,0.00076845585,0.002755386,0.0041707237,0.0019441692,0.0028835547],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011063298,0.0011356584,0.1433474,0.0039348667,0.0010726532,0.0020138717,0.0019773168,0.015973434,0.0064951372,0.0085618105,0.65904856,0.15533294],"study_design_scores_gemma":[0.0005355195,0.0004015341,0.13831869,0.00088152743,0.0004961638,0.0032280777,0.0025469435,0.08962923,0.020934688,0.01536166,0.7272929,0.00037318107],"about_ca_topic_score_codex":0.014654156,"about_ca_topic_score_gemma":0.015944391,"teacher_disagreement_score":0.995552,"about_ca_system_score_codex":0.0011403164,"about_ca_system_score_gemma":0.0024729806,"threshold_uncertainty_score":0.02913773},"labels":[],"label_agreement":null},{"id":"W2602804099","doi":"10.1145/3052973.3052974","title":"BinSequence","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":71,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada; Ministère de la Défense Nationale","keywords":"Computer science; Longest common subsequence problem; Compiler; Block (permutation group theory); Key (lock); Code (set theory); Reuse; Source code; Programming language; Process (computing); Theoretical computer science; Data mining; Algorithm","score_opus":0.03803123773101391,"score_gpt":0.308092317794876,"score_spread":0.2700610800638621,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2602804099","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08202818,0.0011465459,0.6443835,0.0003339609,0.00039441918,0.00077195547,0.017679222,0.23716933,0.01609287],"genre_scores_gemma":[0.2705104,0.0003621142,0.64932305,0.00034734595,0.00008540905,0.0007447059,0.04705443,0.015796615,0.015775843],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982161,0.00014800332,0.00017351314,0.00051099877,0.0008205804,0.00013083739],"domain_scores_gemma":[0.9964981,0.0011993544,0.00039920534,0.0011027119,0.0006850401,0.00011548537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011880848,0.0014462653,0.0009946267,0.005675505,0.00088181545,0.0014208372,0.0021212779,0.0010008033,0.018560948],"category_scores_gemma":[0.009574952,0.00075522816,0.0011842117,0.0033550141,0.0006281584,0.0028105874,0.002370395,0.0007650287,0.007313941],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001358806,0.00038497633,0.016844068,0.0011899443,0.00027519758,0.0005236107,0.00072907656,0.008630004,0.026187534,0.018921267,0.09161603,0.8333396],"study_design_scores_gemma":[0.00034385882,0.0010226374,0.027523808,0.00031127036,0.00020998047,0.0029447323,0.0007103918,0.38644984,0.2188294,0.05367112,0.3076925,0.00029054662],"about_ca_topic_score_codex":0.0020156654,"about_ca_topic_score_gemma":0.003618789,"teacher_disagreement_score":0.018560948,"about_ca_system_score_codex":0.00057439023,"about_ca_system_score_gemma":0.0010371533,"threshold_uncertainty_score":0.062092543},"labels":[],"label_agreement":null},{"id":"W2604153767","doi":"10.1007/s10664-017-9510-8","title":"An empirical study of unspecified dependencies in make-based build systems","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; Polytechnique Montréal; McGill University; Queen's University","funders":"","keywords":"Deliverable; Computer science; Software engineering; Software deployment; sync; Software; Systems engineering; Programming language; Engineering","score_opus":0.05014827432941954,"score_gpt":0.3386470064963655,"score_spread":0.28849873216694594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2604153767","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99773896,0.000019538065,0.0008807025,0.00004083932,0.0000015109692,0.00001340265,0.000030259673,0.000011009263,0.0012637179],"genre_scores_gemma":[0.998966,0.000015672107,0.0006164862,0.000008614969,0.0000014479649,0.0000107544565,0.00006183406,0.000006861847,0.0003123065],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9977513,0.0007990406,0.00015080557,0.00024819048,0.00077637716,0.00027433984],"domain_scores_gemma":[0.867914,0.10186554,0.013511287,0.009867588,0.004939627,0.0019018602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035449464,0.00026431622,0.00020275658,0.0012112268,0.0014893522,0.0011468006,0.0008683111,0.0007290234,0.0033780248],"category_scores_gemma":[0.05806178,0.0004041924,0.00022594762,0.0013648564,0.0015614786,0.0029035306,0.0013569619,0.0015715716,0.00031879087],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070468,0.0029655052,0.89507616,0.0002098973,0.00008975249,0.00074244227,0.01903222,0.010622886,0.0043431944,0.020286635,0.0009216318,0.045004994],"study_design_scores_gemma":[0.000092074886,0.0010725505,0.9195574,0.000103215796,0.000106780324,0.0007330982,0.01630853,0.038312912,0.0060347086,0.01184468,0.0057740286,0.000060067774],"about_ca_topic_score_codex":0.005534154,"about_ca_topic_score_gemma":0.012360718,"teacher_disagreement_score":0.005534154,"about_ca_system_score_codex":0.0013851521,"about_ca_system_score_gemma":0.0017186211,"threshold_uncertainty_score":0.018747687},"labels":[],"label_agreement":null},{"id":"W2604364937","doi":"10.1016/j.entcom.2017.04.001","title":"Open source computer game application: An empirical analysis of quality concerns","year":2017,"lang":"en","type":"article","venue":"Entertainment Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University","funders":"","keywords":"Maintainability; Computer science; Quality (philosophy); Correctness; Software quality; Usability; Reliability (semiconductor); Popularity; Empirical research; Computer game; Source code; Software; Software engineering; Multimedia; Human–computer interaction; Software development","score_opus":0.07709786021378572,"score_gpt":0.4222264552691271,"score_spread":0.3451285950553414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2604364937","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9964658,0.0000823292,0.0009724918,0.00009200573,0.000004960391,0.00008035872,0.00006363722,0.000026911539,0.0022113153],"genre_scores_gemma":[0.9984561,0.000048620957,0.0007057553,0.000042099487,0.000005414982,0.000037153517,0.00012921928,0.000026960199,0.0005487404],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9939084,0.0019823622,0.00039155158,0.00040819467,0.003017645,0.00029180566],"domain_scores_gemma":[0.81533027,0.13298354,0.023029365,0.0038580853,0.020798353,0.0040003755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006985881,0.00035251505,0.00027985882,0.0028845228,0.0008434147,0.002621679,0.00091323,0.0008187976,0.0018134784],"category_scores_gemma":[0.11602606,0.00028865488,0.000281514,0.002259673,0.0010161409,0.0025672452,0.0014513189,0.001544577,0.0003901469],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014166839,0.005395106,0.90878874,0.00042456397,0.00017253475,0.00054899085,0.015940605,0.0007423052,0.005003345,0.0029540733,0.0018894334,0.056723736],"study_design_scores_gemma":[0.00008180366,0.0013787659,0.9652222,0.00014287398,0.00017194804,0.0007374311,0.0114539685,0.012001132,0.0032710203,0.0012597888,0.0042167585,0.000062307],"about_ca_topic_score_codex":0.0044513876,"about_ca_topic_score_gemma":0.005976111,"teacher_disagreement_score":0.006985881,"about_ca_system_score_codex":0.0013377023,"about_ca_system_score_gemma":0.0013088002,"threshold_uncertainty_score":0.036945343},"labels":[],"label_agreement":null},{"id":"W2604385322","doi":"10.1145/3036290.3036323","title":"Investigating the Accuracy of Test Code Size Prediction using Use Case Metrics and Machine Learning Algorithms","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Computer science; Machine learning; Regression testing; Software quality; Software metric; Artificial intelligence; Naive Bayes classifier; Algorithm; Test case; Data mining; Software; Software regression; Code coverage; Regression analysis; Software system; Software development; Support vector machine; Software construction; Programming language","score_opus":0.0977008421245475,"score_gpt":0.32379953975128817,"score_spread":0.22609869762674067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2604385322","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9567105,0.0007669235,0.039542716,0.00024377297,0.000037198613,0.0000650909,0.00066855707,0.0007711215,0.0011940965],"genre_scores_gemma":[0.9808653,0.00014070008,0.01747641,0.000019138248,0.000019052439,0.000045221997,0.0011704545,0.00003447346,0.00022923239],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98904735,0.0049775713,0.0011350719,0.0015291274,0.0028371343,0.0004737996],"domain_scores_gemma":[0.779026,0.19009672,0.010483049,0.0070068743,0.012238575,0.0011488353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016299466,0.0016471589,0.001049518,0.006526669,0.00038466608,0.0015034033,0.0018420087,0.0019378435,0.00061581866],"category_scores_gemma":[0.11063329,0.00046217276,0.0010165462,0.0038832074,0.00056973624,0.0029337464,0.00071685476,0.0012742316,0.00042533688],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006371079,0.0008363212,0.5238229,0.0003264113,0.00058489287,0.000160656,0.00029670724,0.30785075,0.0021171642,0.00073552795,0.0012970119,0.16133456],"study_design_scores_gemma":[0.000013129267,0.00023210788,0.045595784,0.000023694607,0.00003823979,0.000062006235,0.000058797545,0.9515557,0.0017550382,0.0004152561,0.00023182161,0.000018346533],"about_ca_topic_score_codex":0.010139171,"about_ca_topic_score_gemma":0.0069672256,"teacher_disagreement_score":0.016299466,"about_ca_system_score_codex":0.0011064985,"about_ca_system_score_gemma":0.00094373163,"threshold_uncertainty_score":0.08620089},"labels":[],"label_agreement":null},{"id":"W2604794021","doi":"10.1007/s10664-017-9514-4","title":"What do developers search for on the web?","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":177,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of British Columbia","funders":"Ministry of Science and Technology of the People's Republic of China; National Natural Science Foundation of China; Baidu","keywords":"Computer science; Debugging; World Wide Web; Reuse; Software bug; Software; Search engine optimization; Application programming interface; Information retrieval; Software engineering; Search engine; Data science; Programming language; Engineering","score_opus":0.06330645191543763,"score_gpt":0.3294800053743171,"score_spread":0.2661735534588795,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2604794021","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90157074,0.006162768,0.003219423,0.020083899,0.00013232313,0.00008436178,0.0010704113,0.00023261647,0.06744352],"genre_scores_gemma":[0.9887983,0.0018903362,0.0015675548,0.000937304,0.00008712242,0.000021050539,0.0004256868,0.00013696482,0.006135672],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9953418,0.0018379434,0.00025831797,0.0004364643,0.0015863832,0.0005390823],"domain_scores_gemma":[0.9343048,0.041046027,0.01022754,0.0025682992,0.009015879,0.0028375091],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039094575,0.00033830228,0.0005178504,0.004285086,0.0011906287,0.0049642725,0.0007569347,0.002016313,0.007689649],"category_scores_gemma":[0.0618526,0.00035187794,0.00026616175,0.004783501,0.0009729045,0.009995463,0.0010017176,0.0012318451,0.0026782625],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019586203,0.00050033536,0.7339058,0.0008222071,0.00018005597,0.0011457843,0.0145411035,0.00040753445,0.0012729559,0.011933501,0.03618853,0.19890633],"study_design_scores_gemma":[0.00018233275,0.0002790456,0.7477002,0.0022317444,0.00049175555,0.0053750966,0.093205616,0.008510273,0.0047867727,0.03491624,0.102153875,0.00016694353],"about_ca_topic_score_codex":0.012568949,"about_ca_topic_score_gemma":0.020162858,"teacher_disagreement_score":0.012568949,"about_ca_system_score_codex":0.0012720185,"about_ca_system_score_gemma":0.0026512344,"threshold_uncertainty_score":0.02572447},"labels":[],"label_agreement":null},{"id":"W2605594016","doi":"10.1007/s10664-017-9516-2","title":"Data Transformation in Cross-project Defect Prediction","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Transformation (genetics); Computer science; Data mining; Data transformation; Software; Rank (graph theory); Predictive modelling; Reliability engineering; Machine learning; Data warehouse; Mathematics; Engineering","score_opus":0.08894280679107658,"score_gpt":0.37317781031934233,"score_spread":0.28423500352826575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2605594016","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50299144,0.0005153908,0.47122696,0.00078314263,0.00024101668,0.00078550103,0.008864594,0.009648155,0.0049437406],"genre_scores_gemma":[0.8421451,0.0001157586,0.14681044,0.00008583666,0.000033772405,0.00047606972,0.008271557,0.0004577289,0.001603731],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98544276,0.0088652335,0.0014171392,0.0017161733,0.0021169046,0.00044187243],"domain_scores_gemma":[0.92028856,0.056201637,0.0026348708,0.015187396,0.0052143284,0.0004731547],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012857264,0.00066875207,0.0007455603,0.003691411,0.0006273522,0.0018912931,0.0012301918,0.00083883107,0.0038284208],"category_scores_gemma":[0.087343566,0.0004444223,0.0011693521,0.005860583,0.00073905656,0.002209198,0.0020011885,0.0016991056,0.0027729508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002275468,0.0008916195,0.20126693,0.00041516073,0.00035385563,0.00038706756,0.000808296,0.018975893,0.0072687143,0.006657164,0.009652003,0.75104785],"study_design_scores_gemma":[0.00038416547,0.0014722005,0.18366134,0.0003631324,0.00046238198,0.0017323985,0.0026289022,0.68349946,0.05852988,0.03975395,0.027351748,0.00016043105],"about_ca_topic_score_codex":0.00304235,"about_ca_topic_score_gemma":0.002480721,"teacher_disagreement_score":0.012857264,"about_ca_system_score_codex":0.00047276204,"about_ca_system_score_gemma":0.0015543049,"threshold_uncertainty_score":0.0679965},"labels":[],"label_agreement":null},{"id":"W2605988416","doi":"10.7287/peerj.preprints.2617v1","title":"Curating GitHub for engineered software projects","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Classifier (UML); Software engineering; Set (abstract data type); Software bug; Software metric; Skew; Software development; Data mining; Machine learning; Data science; Software quality; Artificial intelligence; Programming language","score_opus":0.030838968956762355,"score_gpt":0.2727086505746935,"score_spread":0.24186968161793115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2605988416","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.72963697,0.00470881,0.1309038,0.0023647437,0.00049387,0.002762731,0.07653728,0.034423314,0.018168464],"genre_scores_gemma":[0.47917128,0.0017380932,0.36834854,0.00038651653,0.00016807744,0.0018139834,0.13869962,0.0027562606,0.0069176345],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9912901,0.0014314743,0.0013589048,0.0020822673,0.0033109838,0.0005262309],"domain_scores_gemma":[0.9504173,0.011202138,0.0144771505,0.011353041,0.010683198,0.0018672642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077548143,0.0013460367,0.0007814755,0.024407767,0.001587618,0.0031453301,0.0017727561,0.0011525155,0.002047325],"category_scores_gemma":[0.053814694,0.0006623246,0.0009609798,0.016224753,0.0011447961,0.0033185773,0.006240573,0.0014492122,0.0023525637],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047560164,0.00025150986,0.4143955,0.003244899,0.00030990606,0.0021701285,0.0104253935,0.003960328,0.01234407,0.008884588,0.11496254,0.4285756],"study_design_scores_gemma":[0.00009770742,0.00029637266,0.63351446,0.0012292508,0.00022274241,0.0025649224,0.007176461,0.046892934,0.0298673,0.011289372,0.2665604,0.00028808464],"about_ca_topic_score_codex":0.010208414,"about_ca_topic_score_gemma":0.02341218,"teacher_disagreement_score":0.024407767,"about_ca_system_score_codex":0.0015071637,"about_ca_system_score_gemma":0.0036768927,"threshold_uncertainty_score":0.04101187},"labels":[],"label_agreement":null},{"id":"W2606008707","doi":"10.71781/11096","title":"Visualisation de la qualité des logiciels de grandes taille","year":2006,"lang":"fr","type":"dissertation","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Humanities; Mathematics; Art","score_opus":0.07566068729090801,"score_gpt":0.4030406045720971,"score_spread":0.3273799172811891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2606008707","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37528643,0.016539479,0.397435,0.006638822,0.0010609383,0.00023743989,0.040373225,0.06467729,0.097751446],"genre_scores_gemma":[0.72575194,0.005743957,0.19802289,0.00032676107,0.00021184035,0.00027141444,0.01386389,0.005231374,0.050576035],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989988,0.00012421905,0.00005714695,0.00019250052,0.0005045734,0.00012271569],"domain_scores_gemma":[0.9932724,0.0027391212,0.0004750496,0.0003648896,0.002814698,0.00033387658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016415062,0.0011565842,0.0009207763,0.009841507,0.0012247663,0.006268821,0.00079921493,0.0013228625,0.024174482],"category_scores_gemma":[0.010163967,0.00068517716,0.00082094886,0.006010255,0.0006048633,0.003408641,0.001422059,0.0016982494,0.004688729],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001631085,0.00020430397,0.037520222,0.0049080984,0.0006221935,0.0007990849,0.0070293965,0.0373104,0.1373846,0.05251333,0.09309637,0.6269809],"study_design_scores_gemma":[0.00038005447,0.00034972778,0.21120502,0.0016952578,0.0005799531,0.001288338,0.0046914755,0.23759905,0.12296439,0.04034189,0.3783691,0.00053571357],"about_ca_topic_score_codex":0.07825818,"about_ca_topic_score_gemma":0.08467335,"teacher_disagreement_score":0.07825818,"about_ca_system_score_codex":0.0022995982,"about_ca_system_score_gemma":0.0025534409,"threshold_uncertainty_score":0.15560538},"labels":[],"label_agreement":null},{"id":"W2606075822","doi":"10.1109/icpc.2017.1","title":"Studying the Prevalence of Exception Handling Anti-Patterns","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Exception handling; Java; Computer science; Crash; Program comprehension; Software; Programming language; Software quality; Source code; Software bug; Software engineering; Software system; Software development","score_opus":0.05192843582252474,"score_gpt":0.3142052720638973,"score_spread":0.26227683624137255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2606075822","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9938378,0.0005485109,0.0036735435,0.00013590862,0.000014109896,0.00003316081,0.0002585135,0.00009899412,0.0013994855],"genre_scores_gemma":[0.9938797,0.00042957772,0.004105505,0.00008664655,0.000024528325,0.000054728393,0.0007171919,0.00006733967,0.0006348505],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99019355,0.0022119205,0.001839688,0.0019729102,0.0031456624,0.0006362759],"domain_scores_gemma":[0.8503572,0.07725492,0.03829142,0.009011095,0.022849644,0.002235764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006044708,0.00039478688,0.00043505765,0.0062226662,0.0007055454,0.0016473548,0.0009456186,0.0007881224,0.0008479091],"category_scores_gemma":[0.0539925,0.0003122265,0.0004353817,0.003581279,0.0010453749,0.0031696665,0.0021217896,0.00090702757,0.00033310754],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012368066,0.000109633314,0.93012846,0.00037031466,0.00008742564,0.00029900152,0.008010492,0.000229474,0.004208061,0.00078367506,0.0005600403,0.055089746],"study_design_scores_gemma":[0.000011789377,0.00026600691,0.96560174,0.0003344182,0.0001779346,0.0022134406,0.009929938,0.0039097993,0.0062940815,0.0019993167,0.009175519,0.000085927604],"about_ca_topic_score_codex":0.0016786635,"about_ca_topic_score_gemma":0.0024411343,"teacher_disagreement_score":0.0062226662,"about_ca_system_score_codex":0.0004071737,"about_ca_system_score_gemma":0.00096487807,"threshold_uncertainty_score":0.03196782},"labels":[],"label_agreement":null},{"id":"W2606446052","doi":"10.1007/978-3-319-57735-7_8","title":"How are Developers Treating License Inconsistency Issues? A Case Study on License Inconsistency Evolution in FOSS Projects","year":2017,"lang":"en","type":"book-chapter","venue":"IFIP advances in information and communication technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science","keywords":"License; Computer science; MIT License; Open source; Operating system; Software","score_opus":0.02581282395413771,"score_gpt":0.3007542597112004,"score_spread":0.2749414357570627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2606446052","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98436594,0.00031002975,0.00909507,0.0013007682,0.000021765367,0.00007833676,0.000047923393,0.00008568016,0.004694474],"genre_scores_gemma":[0.98959875,0.00023198826,0.008357196,0.00013921545,0.000013682225,0.000033911598,0.00006934102,0.00006422807,0.001491606],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9868665,0.007065573,0.00076897256,0.00096657814,0.0034877528,0.0008446204],"domain_scores_gemma":[0.9239441,0.0504672,0.010421473,0.004226168,0.009023265,0.0019177337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013034857,0.00040801475,0.00036329537,0.0031877856,0.0044487067,0.0033922393,0.001787523,0.002548891,0.0015078959],"category_scores_gemma":[0.053181764,0.00054953445,0.00043370572,0.0032769828,0.002937117,0.0057381494,0.0032561352,0.002123239,0.00029837206],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044887277,0.0013358986,0.4446033,0.00060129946,0.00010552082,0.05673956,0.2554035,0.006713235,0.013877497,0.015568726,0.0052951663,0.19930743],"study_design_scores_gemma":[0.00013227508,0.0010501684,0.290613,0.0013532068,0.0003129118,0.043314885,0.4568015,0.058020968,0.036562424,0.026049502,0.08538231,0.00040682848],"about_ca_topic_score_codex":0.006983989,"about_ca_topic_score_gemma":0.01002371,"teacher_disagreement_score":0.013034857,"about_ca_system_score_codex":0.00289167,"about_ca_system_score_gemma":0.0024624413,"threshold_uncertainty_score":0.06893569},"labels":[],"label_agreement":null},{"id":"W2607752446","doi":"","title":"Towards a Framework for Software Product Maturity Measurement","year":2015,"lang":"en","type":"article","venue":"International Conference on Software Engineering Advances","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Capability Maturity Model; Maturity (psychological); Computer science; Product (mathematics); Software measurement; Software; Software engineering; Software quality; Software development; Mathematics; Operating system; Psychology","score_opus":0.08699938751624543,"score_gpt":0.3322893598957313,"score_spread":0.24528997237948585,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2607752446","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024481462,0.0001768303,0.99280375,0.0004493804,0.00002810771,0.00025503058,0.00017462655,0.0019408403,0.0017233413],"genre_scores_gemma":[0.04362697,0.00016205008,0.954312,0.00009321598,0.00002631869,0.0004325079,0.00051979284,0.00015355679,0.0006736573],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97617835,0.007245556,0.0044514337,0.0029603408,0.008296355,0.00086796295],"domain_scores_gemma":[0.97187257,0.007844812,0.003179818,0.004592924,0.011464267,0.0010455985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020434763,0.0022827466,0.0019010893,0.011405321,0.002096602,0.013049993,0.004264422,0.0036796606,0.0021711995],"category_scores_gemma":[0.04578292,0.0016806081,0.002473208,0.0066068685,0.0026124127,0.012407765,0.0071508754,0.0055812164,0.0024488287],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000106622974,0.000605895,0.019420262,0.0007547464,0.0003158473,0.00030523507,0.002376901,0.02155174,0.01136659,0.48102954,0.008516455,0.45365015],"study_design_scores_gemma":[0.000058631453,0.00038130733,0.011175097,0.0012445828,0.00034233998,0.0006415267,0.0019430204,0.47381917,0.013610407,0.42571458,0.07080893,0.00026040873],"about_ca_topic_score_codex":0.010792689,"about_ca_topic_score_gemma":0.010301375,"teacher_disagreement_score":0.020434763,"about_ca_system_score_codex":0.0037319527,"about_ca_system_score_gemma":0.007282287,"threshold_uncertainty_score":0.10807067},"labels":[],"label_agreement":null},{"id":"W2608234913","doi":"10.1109/pst.2016.7906955","title":"Analysing vulnerability reproducibility for Firefox browser","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Vulnerability (computing); Random forest; Decision tree; Software; Machine learning; Naive Bayes classifier; Data mining; Software security assurance; Logistic regression; Asset (computer security); Computer security; Artificial intelligence; Information security; Support vector machine","score_opus":0.041091559843379595,"score_gpt":0.31716460246613787,"score_spread":0.2760730426227583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2608234913","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9951167,0.00013690724,0.003245343,0.000020243353,0.0000043204914,0.000025261852,0.0004656666,0.0006034219,0.00038225777],"genre_scores_gemma":[0.99312055,0.00006244696,0.0052037057,0.0000050370304,0.0000056521753,0.000021689479,0.0011966478,0.00009717693,0.0002871174],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.998198,0.00038520948,0.0001966679,0.00037366006,0.00067404285,0.00017251521],"domain_scores_gemma":[0.9593208,0.026487313,0.0065409443,0.003055163,0.0038697785,0.0007260202],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0034603325,0.00055262056,0.00031184067,0.0056026806,0.0003444566,0.00067291374,0.00040655083,0.00045235726,0.0005303985],"category_scores_gemma":[0.022570869,0.0002232409,0.0007262207,0.002098789,0.0004947164,0.0012692255,0.00062050426,0.0006175772,0.00022826169],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000435299,0.0002710944,0.8658579,0.00027042022,0.000342172,0.0008024724,0.0031052434,0.017637886,0.01566113,0.0007698818,0.001975845,0.092870794],"study_design_scores_gemma":[0.000012593042,0.00039422885,0.89360636,0.000060485643,0.00011719825,0.0010854707,0.00046308246,0.09380634,0.008674892,0.00052165636,0.0011956261,0.00006204078],"about_ca_topic_score_codex":0.010367286,"about_ca_topic_score_gemma":0.010056863,"teacher_disagreement_score":0.99653965,"about_ca_system_score_codex":0.00068568473,"about_ca_system_score_gemma":0.00050924753,"threshold_uncertainty_score":0.02061385},"labels":[],"label_agreement":null},{"id":"W2608917368","doi":"10.4236/jsea.2017.104020","title":"Estimation Models for Software Functional Test Effort","year":2017,"lang":"en","type":"article","venue":"Journal of Software Engineering and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Function point; Use Case Points; Benchmarking; Computer science; Software; Estimation; Function (biology); Reliability engineering; Data mining; Software development; Engineering; Software development process; Systems engineering","score_opus":0.022082631807930256,"score_gpt":0.26443589080083596,"score_spread":0.24235325899290572,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2608917368","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15323159,0.0004835732,0.84168786,0.0003441654,0.000025414913,0.00034395134,0.0015056862,0.0008721364,0.0015055706],"genre_scores_gemma":[0.8038555,0.0007559877,0.18401416,0.00008460595,0.000032998716,0.0019817508,0.006090928,0.00023872226,0.0029453537],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99120694,0.005802386,0.00045225065,0.0010669249,0.0009581449,0.0005134308],"domain_scores_gemma":[0.91015357,0.08093862,0.0034308871,0.002071963,0.0032066638,0.00019833486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021697419,0.0018616688,0.0014939616,0.004379878,0.00038044955,0.0023092907,0.0022399717,0.0018366632,0.0022752795],"category_scores_gemma":[0.07347099,0.0010565645,0.0029804178,0.002898827,0.00089948846,0.0027051806,0.0014792687,0.0026090182,0.000885899],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000097495926,0.00008606259,0.019563615,0.00011106058,0.00023103961,0.00005727124,0.0002112288,0.93390673,0.00042929925,0.010016614,0.0007318983,0.034557793],"study_design_scores_gemma":[0.000008491614,0.000046068817,0.0038514617,0.00003665533,0.00003657306,0.000022857368,0.000048392274,0.99134904,0.00024215368,0.004028087,0.0003130962,0.00001709815],"about_ca_topic_score_codex":0.011420929,"about_ca_topic_score_gemma":0.0063341176,"teacher_disagreement_score":0.021697419,"about_ca_system_score_codex":0.0020092183,"about_ca_system_score_gemma":0.0013298133,"threshold_uncertainty_score":0.1147483},"labels":[],"label_agreement":null},{"id":"W2611052873","doi":"10.1109/acit-csii-bcd.2016.042","title":"Using Software Metrics Thresholds to Predict Fault-Prone Classes in Object-Oriented Software","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Computer science; Software; Software fault tolerance; Software metric; Software sizing; Software measurement; Software quality; Software construction; Software development; Software engineering; Programming language","score_opus":0.03900852194745664,"score_gpt":0.3002599662741826,"score_spread":0.261251444326726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2611052873","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9122239,0.0014215778,0.080350526,0.00018399293,0.000054879954,0.00008640196,0.0019859415,0.0018892054,0.0018035352],"genre_scores_gemma":[0.97315514,0.00018788702,0.02400515,0.000018803463,0.00001796093,0.000047712867,0.0022817764,0.000071599956,0.00021395214],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99634105,0.00079422153,0.0005070043,0.0007466701,0.0013596051,0.0002513422],"domain_scores_gemma":[0.95894486,0.028269615,0.005325065,0.0023176493,0.004424159,0.00071871624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058843317,0.0011296705,0.0005849434,0.008805272,0.0003102607,0.0012679985,0.00057610305,0.0010170445,0.00041177115],"category_scores_gemma":[0.036673695,0.00023398497,0.00063390704,0.0035477525,0.00035245408,0.0019119575,0.0006962642,0.00074693136,0.0004115818],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006267991,0.0003147045,0.66557765,0.00035864726,0.00033229563,0.00014001105,0.00042368256,0.0898108,0.0059458497,0.0013342681,0.0029502662,0.23218499],"study_design_scores_gemma":[0.00004623134,0.00064413185,0.34571674,0.00016372126,0.00015497077,0.00047808798,0.00034646437,0.62755305,0.015749158,0.0064956406,0.002538517,0.0001131903],"about_ca_topic_score_codex":0.0036889638,"about_ca_topic_score_gemma":0.0041446197,"teacher_disagreement_score":0.008805272,"about_ca_system_score_codex":0.0007613646,"about_ca_system_score_gemma":0.00057413295,"threshold_uncertainty_score":0.031119704},"labels":[],"label_agreement":null},{"id":"W2612705982","doi":"10.1007/s10664-017-9522-4","title":"Identifying self-admitted technical debt in open source projects using text mining","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":190,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Concordia University","funders":"Ministry of Science and Technology of the People's Republic of China; National Natural Science Foundation of China","keywords":"Technical debt; Computer science; Classifier (UML); Source code; Code review; Baseline (sea); Open source; Artificial intelligence; F1 score; Machine learning; Data mining; Natural language processing; Software; Software quality; Software development; Programming language","score_opus":0.07815053665037296,"score_gpt":0.3568263279761991,"score_spread":0.27867579132582615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2612705982","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9961922,0.00013924683,0.0013053943,0.00016543905,0.000012149558,0.000024086452,0.0011683499,0.000032995325,0.00096015324],"genre_scores_gemma":[0.9939329,0.00012884014,0.0021977734,0.00004973359,0.00003499869,0.00004894513,0.0028231137,0.000016708245,0.0007670712],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9970227,0.00065216847,0.0007276429,0.00044132094,0.0009093334,0.0002468158],"domain_scores_gemma":[0.9268764,0.038814265,0.022415081,0.0024538063,0.007328088,0.002112351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034723664,0.00026334552,0.00030170282,0.0068802545,0.000713843,0.0016122028,0.0006649135,0.00096129486,0.001028692],"category_scores_gemma":[0.03971322,0.00018601121,0.0002806805,0.005615596,0.00039532126,0.0021524664,0.001209338,0.00082593726,0.0005037529],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001217976,0.00020806324,0.96066344,0.000117136304,0.00004186226,0.00030680888,0.0011632809,0.0003719807,0.0019374135,0.00049447996,0.0014063134,0.03316738],"study_design_scores_gemma":[0.000014233797,0.00011045082,0.97472715,0.0001559473,0.000056431916,0.0005643507,0.0032981082,0.012678953,0.002632536,0.002212856,0.0035181197,0.000030956366],"about_ca_topic_score_codex":0.0023751673,"about_ca_topic_score_gemma":0.004420753,"teacher_disagreement_score":0.0068802545,"about_ca_system_score_codex":0.0005527096,"about_ca_system_score_gemma":0.0009208096,"threshold_uncertainty_score":0.018363893},"labels":[],"label_agreement":null},{"id":"W2612865512","doi":"10.1007/s10664-017-9520-6","title":"An empirical study of the integration time of fixed issues","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; McGill University","funders":"","keywords":"Eclipse; Computer science; Fixed cost; Heuristics; Fixed point; Fixed effects model; Operations research; Engineering; Business; Mathematics; Accounting; Operating system; Statistics","score_opus":0.038382145223893084,"score_gpt":0.3498034475928034,"score_spread":0.3114213023689103,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2612865512","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9902497,0.00028535357,0.0027485495,0.00014709546,0.000010486279,0.00003673552,0.00007503449,0.000025885469,0.0064212186],"genre_scores_gemma":[0.9976259,0.00008279391,0.0012671261,0.00002469106,0.000013011212,0.00001561402,0.00010471432,0.00001766712,0.00084852584],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9966343,0.0017143543,0.00022449088,0.00032789007,0.00082121143,0.00027771253],"domain_scores_gemma":[0.6979755,0.24966656,0.031495687,0.0095408065,0.007478574,0.0038429582],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006418106,0.00025319742,0.00028685026,0.0017463124,0.0007926894,0.0022902372,0.00080839975,0.00092066085,0.011582374],"category_scores_gemma":[0.1561751,0.0003197567,0.0002998201,0.002821774,0.0010702132,0.003943597,0.0011473518,0.0027008604,0.0007539744],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003859422,0.006727056,0.7013536,0.00042172044,0.00022034314,0.0008568135,0.013396658,0.0068517854,0.008878953,0.04480534,0.0028065713,0.20982176],"study_design_scores_gemma":[0.00049988466,0.0046566417,0.8676884,0.00023039484,0.00042879867,0.0013994046,0.014606543,0.055430934,0.008212146,0.032576714,0.014144643,0.00012545657],"about_ca_topic_score_codex":0.0018088244,"about_ca_topic_score_gemma":0.0013535467,"teacher_disagreement_score":0.011582374,"about_ca_system_score_codex":0.001006482,"about_ca_system_score_gemma":0.0010141858,"threshold_uncertainty_score":0.038746893},"labels":[],"label_agreement":null},{"id":"W26133019","doi":"10.1097/ede.0000000000000319","title":"Using Clustering Technique to Restructure Programs.","year":2004,"lang":"en","type":"article","venue":"Software Engineering Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Restructuring; Cohesion (chemistry); Code refactoring; Computer science; Group cohesiveness; Business process reengineering; Cluster analysis; Measure (data warehouse); Software; Function (biology); Software engineering; Industrial engineering; Programming language; Database; Operations management; Business; Artificial intelligence; Engineering","score_opus":0.08210969072528825,"score_gpt":0.3807452638190845,"score_spread":0.2986355730937963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W26133019","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016455302,0.00018676759,0.94865745,0.00039768414,0.00024019198,0.0008661912,0.003182189,0.019634902,0.010379328],"genre_scores_gemma":[0.042343907,0.00012209539,0.9402304,0.00010776545,0.000037433387,0.0008722606,0.004968816,0.002268648,0.0090485895],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99777836,0.00044775265,0.0001991335,0.0006795492,0.0006756139,0.00021958751],"domain_scores_gemma":[0.9946767,0.0012728205,0.00044812317,0.0018160554,0.0016571714,0.00012908249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021811188,0.0011619455,0.0009561759,0.005610244,0.0018080488,0.0018579882,0.0020202075,0.0011257918,0.021035293],"category_scores_gemma":[0.014158594,0.0008239151,0.0018773404,0.0065966086,0.0006596976,0.0014851948,0.001954663,0.0022754385,0.010520307],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001712048,0.00027037624,0.00599606,0.0003955694,0.00017170778,0.00014532736,0.0011909916,0.038643833,0.010552597,0.020518042,0.034825385,0.8871188],"study_design_scores_gemma":[0.00010244148,0.00026990715,0.017350392,0.00023714415,0.00025311264,0.0005540594,0.0013891278,0.6303533,0.0441847,0.070803404,0.23426242,0.00024009106],"about_ca_topic_score_codex":0.024376534,"about_ca_topic_score_gemma":0.029770372,"teacher_disagreement_score":0.024376534,"about_ca_system_score_codex":0.0015827616,"about_ca_system_score_gemma":0.0031022795,"threshold_uncertainty_score":0.07037002},"labels":[],"label_agreement":null},{"id":"W2614198659","doi":"10.1109/icsm.2004.1357817","title":"An empirical study of .ne-grained software modi .cations","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Visualization; Software visualization; Software; Software evolution; Empirical research; Software engineering; Software system; Programming language; Software construction; Data mining","score_opus":0.03642011555274559,"score_gpt":0.34479311202092344,"score_spread":0.30837299646817784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2614198659","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99820936,0.00008290793,0.00060948497,0.00010447753,0.0000031878083,0.000042840984,0.00022823166,0.000014387754,0.0007051026],"genre_scores_gemma":[0.99791664,0.00008190737,0.0012030015,0.000033082113,0.000007258089,0.00006414328,0.00037528083,0.000008886752,0.00030980783],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98970187,0.0050160736,0.0012591708,0.0012133667,0.0023519585,0.00045749027],"domain_scores_gemma":[0.5581212,0.34401935,0.062239684,0.013920037,0.017592717,0.004106978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010377452,0.0003164007,0.00030458163,0.0026315588,0.0009064215,0.0012684707,0.0010992299,0.00097496394,0.0019552025],"category_scores_gemma":[0.14072818,0.000413716,0.00024849648,0.0044416613,0.0016411485,0.004075031,0.0010557626,0.0015676633,0.0005033991],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037511418,0.0014981361,0.9642683,0.00024481764,0.00008761718,0.00036392658,0.0063601113,0.0009429065,0.0011684486,0.0010727162,0.0010432735,0.022574672],"study_design_scores_gemma":[0.000060064576,0.0011426955,0.9742716,0.00006314718,0.0000460691,0.0009129432,0.009837982,0.0071411696,0.0015069239,0.0008330346,0.0041400846,0.000044204364],"about_ca_topic_score_codex":0.0028805411,"about_ca_topic_score_gemma":0.0049582724,"teacher_disagreement_score":0.010377452,"about_ca_system_score_codex":0.0008626037,"about_ca_system_score_gemma":0.000921838,"threshold_uncertainty_score":0.05488187},"labels":[],"label_agreement":null},{"id":"W2614313580","doi":"10.1142/s0218843017420011","title":"Semantic Analysis of RESTful APIs for the Detection of Linguistic Patterns and Antipatterns","year":2017,"lang":"en","type":"article","venue":"International Journal of Cooperative Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université du Québec à Montréal; Concordia University","funders":"","keywords":"Computer science; Reusability; Identifier; Documentation; World Wide Web; Software; Software engineering; Programming language","score_opus":0.023432774806279436,"score_gpt":0.31717598576786,"score_spread":0.29374321096158057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2614313580","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30449015,0.00061807263,0.6566637,0.00080078375,0.00009189147,0.0006639104,0.004347222,0.023103597,0.009220675],"genre_scores_gemma":[0.7192499,0.0003198512,0.26844928,0.00030654776,0.000043262884,0.0005058776,0.0070115933,0.00081906584,0.0032945832],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975567,0.00037653744,0.0003708649,0.00031290355,0.0012047141,0.0001781798],"domain_scores_gemma":[0.99361587,0.0017801845,0.0014563026,0.00089841447,0.0021240711,0.00012513538],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013380655,0.00069316145,0.00051265163,0.005354195,0.00081205706,0.0015227397,0.0006500692,0.0005655074,0.0011662719],"category_scores_gemma":[0.006932418,0.0002456977,0.0012752666,0.0022445044,0.000688121,0.0022364552,0.0015940139,0.0007254968,0.0006213031],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052873493,0.0004898503,0.14682357,0.0014450392,0.00030277192,0.0024936222,0.0042724563,0.005612202,0.100851834,0.040547993,0.016873699,0.67975825],"study_design_scores_gemma":[0.000063635496,0.00030121685,0.1387593,0.000454118,0.00045064703,0.0029447624,0.0042720186,0.5729368,0.13907133,0.06642862,0.07406023,0.0002573495],"about_ca_topic_score_codex":0.0036284646,"about_ca_topic_score_gemma":0.00443685,"teacher_disagreement_score":0.005354195,"about_ca_system_score_codex":0.0007903761,"about_ca_system_score_gemma":0.0018350957,"threshold_uncertainty_score":0.0072146654},"labels":[],"label_agreement":null},{"id":"W2615372411","doi":"10.1017/cbo9781139164900","title":"Specifying Software","year":2002,"lang":"en","type":"book","venue":"Cambridge University Press eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Programming language; Compiler; Programmer; Presentation (obstetrics); Software engineering; Software; Notation","score_opus":0.030567518114004134,"score_gpt":0.20712248191079435,"score_spread":0.17655496379679023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2615372411","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011575441,0.010841489,0.7380838,0.0028644414,0.0012343812,0.00037012566,0.001662589,0.0047694547,0.23901616],"genre_scores_gemma":[0.030845283,0.025473053,0.5402196,0.0023448202,0.0009601269,0.0011906157,0.0070068054,0.0034001754,0.38855952],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99770904,0.0005719673,0.0002680516,0.00021401695,0.0011446596,0.000092294285],"domain_scores_gemma":[0.9981437,0.0010843532,0.000103788916,0.00032637973,0.0002833652,0.00005854842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001819543,0.0020884878,0.0008277167,0.0020690223,0.0014104276,0.004800973,0.0014368814,0.0019986816,0.033147056],"category_scores_gemma":[0.003907144,0.0012305178,0.00097247027,0.0028493255,0.0021500578,0.005980443,0.0023698402,0.003154987,0.02428806],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000014338536,0.000022964676,0.00010927753,0.0005241176,0.000013156851,0.00037824345,0.0011503958,0.0015340805,0.0029316426,0.7667336,0.08828734,0.13830091],"study_design_scores_gemma":[0.0000062870245,0.0000120810655,0.00004535849,0.00020513998,0.000004102391,0.0005606421,0.00007851631,0.0012086938,0.0008188347,0.06596342,0.9310822,0.000014763416],"about_ca_topic_score_codex":0.001296794,"about_ca_topic_score_gemma":0.0028268693,"teacher_disagreement_score":0.033147056,"about_ca_system_score_codex":0.0015504069,"about_ca_system_score_gemma":0.0022066904,"threshold_uncertainty_score":0.110888004},"labels":[],"label_agreement":null},{"id":"W2615622384","doi":"10.1109/icst.2017.15","title":"Recovering Semantic Traceability Links between APIs and Security Vulnerabilities: An Ontological Modeling Approach","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Traceability; Software engineering; Secure coding; Ontology; Reuse; Knowledge management; Data science; World Wide Web; Computer security; Software security assurance; Information security; Engineering","score_opus":0.06923665995780864,"score_gpt":0.3121547846778099,"score_spread":0.24291812472000124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2615622384","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022754427,0.00016448429,0.96826106,0.0013591099,0.000033938883,0.00020662062,0.00078904326,0.0005331487,0.005898189],"genre_scores_gemma":[0.23978868,0.000560123,0.75460297,0.00025917872,0.00003379972,0.00042383693,0.0021958707,0.00015212028,0.0019833783],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960743,0.0013511817,0.00046142543,0.0005861753,0.0012556022,0.00027134106],"domain_scores_gemma":[0.99045116,0.0045689647,0.0012538307,0.0021786701,0.0012387848,0.00030865654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055514616,0.00094836415,0.0006271686,0.010130976,0.0026198006,0.0060034674,0.0027369303,0.0023809194,0.0016627705],"category_scores_gemma":[0.015060486,0.0010543836,0.0038697019,0.0063475594,0.0026037092,0.011789232,0.00575365,0.0028430205,0.0004319178],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011298581,0.00055443094,0.015116177,0.00037736213,0.00037771423,0.0020222552,0.006737721,0.17685467,0.005277504,0.6759269,0.0035388346,0.113103405],"study_design_scores_gemma":[0.000028013033,0.00004559522,0.0027252256,0.0002877398,0.00031581952,0.00045157893,0.002449386,0.5843432,0.0037029223,0.37747878,0.028086634,0.00008524996],"about_ca_topic_score_codex":0.03970621,"about_ca_topic_score_gemma":0.037700687,"teacher_disagreement_score":0.03970621,"about_ca_system_score_codex":0.0038073112,"about_ca_system_score_gemma":0.0064715203,"threshold_uncertainty_score":0.07895017},"labels":[],"label_agreement":null},{"id":"W2615656263","doi":"10.1109/icst.2017.27","title":"Broadcast vs. Unicast Review Technology: Does It Matter?","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Unicast; Computer science; Software quality; Quality (philosophy); Software; Code (set theory); Order (exchange); Process (computing); Software engineering; Software development; Computer network; Multicast; Business; Operating system","score_opus":0.02081967236820706,"score_gpt":0.31060405796281365,"score_spread":0.2897843855946066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2615656263","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7313023,0.038664367,0.076234676,0.03558584,0.003443717,0.0055356026,0.0016139423,0.0031627042,0.10445691],"genre_scores_gemma":[0.95572567,0.0050799944,0.026786603,0.0031121452,0.0010862909,0.0018378213,0.0003242917,0.0006339643,0.0054131444],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.7496947,0.14162181,0.02569423,0.012378936,0.06661547,0.003994809],"domain_scores_gemma":[0.22748044,0.5584739,0.10290673,0.02914278,0.07290146,0.009094782],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.121596985,0.0005076027,0.0016584125,0.008013124,0.0024445145,0.011932639,0.0023208782,0.002080911,0.0054867915],"category_scores_gemma":[0.54517823,0.0011747305,0.001557,0.0069274656,0.0030167473,0.013998466,0.004562526,0.00293319,0.0023570892],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004451118,0.0009963169,0.21672165,0.015238793,0.0018805816,0.0006804043,0.035464395,0.0012822996,0.016625956,0.012287558,0.022209942,0.67216104],"study_design_scores_gemma":[0.0027297437,0.01812673,0.6641986,0.016780073,0.006031733,0.006993751,0.0421997,0.0110564325,0.022846712,0.018274847,0.18959218,0.0011695892],"about_ca_topic_score_codex":0.0039065043,"about_ca_topic_score_gemma":0.0046530534,"teacher_disagreement_score":0.121596985,"about_ca_system_score_codex":0.0064752833,"about_ca_system_score_gemma":0.0069388254,"threshold_uncertainty_score":0.64307404},"labels":[],"label_agreement":null},{"id":"W2617519947","doi":"10.1109/icse.2017.31","title":"On Cross-Stack Configuration Errors","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Stack (abstract data type); Call stack; JavaScript; Slicing; Protocol stack; Program slicing; Configuration Management (ITSM); Operating system; Software; Modular design; Source code; Programming language; World Wide Web","score_opus":0.03782559471601254,"score_gpt":0.34553767519679635,"score_spread":0.3077120804807838,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2617519947","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82554436,0.0016848386,0.15235913,0.0014060776,0.00017230547,0.00029041397,0.0007077134,0.0075205266,0.01031463],"genre_scores_gemma":[0.9435967,0.0005238752,0.051652364,0.00030049155,0.000048320017,0.00007984887,0.0006116182,0.0010145683,0.002172148],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9817059,0.0048475526,0.0012852038,0.0024241495,0.008540203,0.0011971024],"domain_scores_gemma":[0.8714523,0.07269017,0.022607403,0.014609491,0.01734992,0.0012908425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077506066,0.0013810345,0.0006233889,0.0066411253,0.0013148548,0.0023433843,0.0013640007,0.0012388823,0.002499772],"category_scores_gemma":[0.0927221,0.0008087588,0.0006446228,0.003913344,0.0019714695,0.0060723913,0.0025018146,0.0016059941,0.0005681127],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008990974,0.00045278948,0.26827025,0.0013106241,0.00028110584,0.0042864853,0.014516157,0.064585604,0.01601929,0.013931229,0.008020323,0.60742706],"study_design_scores_gemma":[0.00011220516,0.001397,0.33450195,0.002654387,0.0008805657,0.009475419,0.019777814,0.42407486,0.10592897,0.050428037,0.0500383,0.00073049654],"about_ca_topic_score_codex":0.00399742,"about_ca_topic_score_gemma":0.004501276,"teacher_disagreement_score":0.0077506066,"about_ca_system_score_codex":0.0016711074,"about_ca_system_score_gemma":0.0019703046,"threshold_uncertainty_score":0.040989578},"labels":[],"label_agreement":null},{"id":"W2620389881","doi":"10.71781/10910","title":"Amélioration de la prédiction de la qualité du logiciel par combinaison et adaptation de modèles","year":2005,"lang":"fr","type":"dissertation","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Humanities; Philosophy","score_opus":0.0578098291164172,"score_gpt":0.3704256713398595,"score_spread":0.3126158422234423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2620389881","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45031604,0.007086999,0.5245118,0.0021975604,0.00049096736,0.0001859959,0.0013578103,0.009382085,0.0044706985],"genre_scores_gemma":[0.83859885,0.0013651965,0.14993253,0.00022501728,0.00011253091,0.00011480445,0.0010981715,0.00032945946,0.008223419],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989458,0.000301263,0.00005587951,0.00030838753,0.00029517416,0.000093551665],"domain_scores_gemma":[0.9950507,0.0031039629,0.0002384936,0.0005615961,0.0009522129,0.00009305085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025347245,0.0014582778,0.001512172,0.0010065646,0.0004006953,0.002121509,0.0011184099,0.0013236507,0.0039624567],"category_scores_gemma":[0.010305863,0.00063186156,0.0012709517,0.0009915531,0.00045372493,0.0022233266,0.0009140735,0.0015483128,0.0009259634],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091134536,0.0005956991,0.022504538,0.0002460871,0.0008473272,0.00013089094,0.00018691704,0.34233016,0.020552147,0.0015227837,0.0050472748,0.6051249],"study_design_scores_gemma":[0.0000365327,0.00015692519,0.006247898,0.00002045304,0.00012551894,0.00003361428,0.000024613873,0.98476267,0.006644444,0.00085389894,0.0010682316,0.000025172787],"about_ca_topic_score_codex":0.06950358,"about_ca_topic_score_gemma":0.06996643,"teacher_disagreement_score":0.06950358,"about_ca_system_score_codex":0.001588749,"about_ca_system_score_gemma":0.0014301989,"threshold_uncertainty_score":0.13819802},"labels":[],"label_agreement":null},{"id":"W2620436109","doi":"10.1109/icse.2017.14","title":"Clone Refactoring with Lambda Expressions","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; Lambda; clone (Java method); Java; Computer science; Programming language; Code (set theory); Software; Artificial intelligence; Biology; Physics; Genetics; Set (abstract data type)","score_opus":0.033987259146658,"score_gpt":0.2982540105640026,"score_spread":0.26426675141734457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2620436109","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42816845,0.0019309382,0.5270916,0.00062910543,0.00020927518,0.0007906238,0.0018227926,0.03450031,0.0048568924],"genre_scores_gemma":[0.44292992,0.000809828,0.53899,0.0004345216,0.00006450981,0.00052128494,0.0048796395,0.005659013,0.0057112486],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9922431,0.0018602152,0.0008865197,0.0015337587,0.0029813226,0.0004949329],"domain_scores_gemma":[0.95907193,0.01564849,0.0040691104,0.014070714,0.0066951257,0.0004445433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008585944,0.0013400071,0.0008911862,0.0026053006,0.00083592744,0.0014317912,0.0021155223,0.0012623558,0.0012615668],"category_scores_gemma":[0.04989153,0.00071988686,0.0017245051,0.0020433404,0.0011809529,0.0022974047,0.0028954481,0.001874206,0.00091298646],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084660214,0.0006608106,0.10656168,0.0020419317,0.00057354046,0.0022582065,0.0052627646,0.05279479,0.15265119,0.013235762,0.012513171,0.6505996],"study_design_scores_gemma":[0.00036790964,0.0015485055,0.059451286,0.0010370439,0.0016442643,0.0036218364,0.0023828987,0.41995236,0.3229306,0.03365618,0.1529211,0.00048604986],"about_ca_topic_score_codex":0.002983826,"about_ca_topic_score_gemma":0.0038713617,"teacher_disagreement_score":0.008585944,"about_ca_system_score_codex":0.0010158907,"about_ca_system_score_gemma":0.002174154,"threshold_uncertainty_score":0.045407355},"labels":[],"label_agreement":null},{"id":"W2620636222","doi":"10.1109/icse-c.2017.3","title":"Fast and Flexible Large-Scale Clone Detection with CloneWorks","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Workstation; Scalability; Flexibility (engineering); Source code; Normalization (sociology); Software maintenance; Software; Software system; Operating system; Mathematics","score_opus":0.011775336914232817,"score_gpt":0.25564508002634645,"score_spread":0.24386974311211362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2620636222","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13014716,0.00076883665,0.6145467,0.00031531067,0.00017466379,0.00044163587,0.0050478145,0.24481155,0.0037463193],"genre_scores_gemma":[0.26214084,0.0002424604,0.71370035,0.00032016129,0.0000501422,0.00087132835,0.011543518,0.007216395,0.00391476],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99591404,0.00053567433,0.00034708803,0.0012425819,0.0016798773,0.0002806647],"domain_scores_gemma":[0.9868314,0.005406464,0.0009998963,0.0041817096,0.0021310127,0.00044937083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029902803,0.0014449836,0.0012417637,0.0035819001,0.0009462914,0.0026858388,0.0027086614,0.0018861432,0.003433025],"category_scores_gemma":[0.0155871445,0.0013018666,0.0012536347,0.0031337661,0.00081094424,0.00419082,0.0031237449,0.0013941056,0.0027442987],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015383941,0.0004677581,0.03076578,0.0008587788,0.00059574493,0.0008845214,0.0014806519,0.013997581,0.30569905,0.007329039,0.06223095,0.5741517],"study_design_scores_gemma":[0.0003773765,0.00084953377,0.0168767,0.00012389198,0.00022318395,0.0015892737,0.0006583814,0.3059148,0.6030084,0.013943674,0.056112025,0.00032273738],"about_ca_topic_score_codex":0.0026192167,"about_ca_topic_score_gemma":0.0032549333,"teacher_disagreement_score":0.0035819001,"about_ca_system_score_codex":0.0009990975,"about_ca_system_score_gemma":0.0011745347,"threshold_uncertainty_score":0.015814304},"labels":[],"label_agreement":null},{"id":"W2620682708","doi":"10.1109/icse-c.2017.11","title":"RACK: Code Search in the IDE Using Crowdsourced Knowledge","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Web search query; Code (set theory); Context (archaeology); Source code; Programming language; Matching (statistics); Query expansion; World Wide Web; Search engine; Database; Set (abstract data type)","score_opus":0.10148763127305913,"score_gpt":0.3820968922068171,"score_spread":0.28060926093375793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2620682708","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07043892,0.0019194883,0.58332026,0.0023082502,0.0004944551,0.0031187823,0.061110854,0.23355444,0.043734487],"genre_scores_gemma":[0.22739358,0.00070221385,0.6811968,0.00070971285,0.00011206385,0.0016897584,0.070160344,0.006610222,0.011425247],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962475,0.0010597707,0.0002830357,0.0010343996,0.0011270292,0.00024826726],"domain_scores_gemma":[0.9909808,0.005074948,0.00048625947,0.0020104628,0.0009551354,0.00049243774],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036737793,0.0017796202,0.0013025792,0.0075584035,0.0013045067,0.0024683,0.0023375012,0.0016949655,0.008388839],"category_scores_gemma":[0.01607011,0.0005512987,0.0012852277,0.0035413625,0.00096251647,0.004532354,0.006507096,0.0011671397,0.00784625],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018681393,0.00091929355,0.010576662,0.0035861633,0.00046318534,0.0018110443,0.004317494,0.018890511,0.023340564,0.020164194,0.26300904,0.65105367],"study_design_scores_gemma":[0.0008632215,0.0005654783,0.010813249,0.0006241576,0.00020628262,0.0010698888,0.00609246,0.53376275,0.041396968,0.09393738,0.31012428,0.0005438883],"about_ca_topic_score_codex":0.009540236,"about_ca_topic_score_gemma":0.015125165,"teacher_disagreement_score":0.009540236,"about_ca_system_score_codex":0.0010473756,"about_ca_system_score_gemma":0.0023297088,"threshold_uncertainty_score":0.028063476},"labels":[],"label_agreement":null},{"id":"W2620701800","doi":"10.1109/icse-c.2017.147","title":"Mining Complex Temporal API Usage Patterns: An Evolutionary Approach","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Task (project management); Software; Genetic programming; Machine learning; Artificial intelligence; Data mining; Programming language","score_opus":0.08968146937094762,"score_gpt":0.3165162103928409,"score_spread":0.22683474102189327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2620701800","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12223742,0.0002340715,0.87428457,0.0008925472,0.000021675149,0.00014554455,0.00018152446,0.0004750626,0.0015275823],"genre_scores_gemma":[0.3323787,0.00026334805,0.66403276,0.00031857126,0.000035666704,0.00022271887,0.000651871,0.00012111289,0.0019752658],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984674,0.0003723799,0.00009997392,0.00049769745,0.00043438017,0.00012818922],"domain_scores_gemma":[0.99548197,0.0026642063,0.0005277886,0.0004190563,0.0007431753,0.00016375759],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024164359,0.00082065124,0.000849096,0.002730037,0.0006942324,0.0013472673,0.0025265191,0.00137664,0.000993416],"category_scores_gemma":[0.00851784,0.0007143173,0.0012482462,0.0022684874,0.0010779179,0.0020234114,0.0014698447,0.0014125509,0.00025788322],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001238979,0.00060547155,0.044106156,0.00026604606,0.00038576324,0.0007381927,0.0012324263,0.36740407,0.0159019,0.018617641,0.0021753605,0.54844314],"study_design_scores_gemma":[0.00001427387,0.00005487013,0.00239619,0.00002142357,0.000051801973,0.00016323246,0.00017538693,0.9825659,0.0016414474,0.011515766,0.00138241,0.000017389953],"about_ca_topic_score_codex":0.00434206,"about_ca_topic_score_gemma":0.007886733,"teacher_disagreement_score":0.00434206,"about_ca_system_score_codex":0.0007807264,"about_ca_system_score_gemma":0.0013961218,"threshold_uncertainty_score":0.012779474},"labels":[],"label_agreement":null},{"id":"W2621103691","doi":"10.1109/icse-c.2017.78","title":"CloneWorks: A Fast and Flexible Large-Scale Near-Miss Clone Detection Tool","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Scalability; Workstation; Source code; Detector; Precision and recall; Code (set theory); Operating system; Artificial intelligence; Programming language; Biology; Gene","score_opus":0.013732591123816052,"score_gpt":0.2698143242420956,"score_spread":0.2560817331182796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2621103691","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039306358,0.0009682641,0.51761264,0.00021809018,0.00015897326,0.00039656102,0.0038883616,0.43406945,0.003381263],"genre_scores_gemma":[0.23536655,0.0005272591,0.71081156,0.00046418892,0.00012662403,0.001042274,0.019014534,0.02345036,0.009196645],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9958259,0.00045301637,0.00031976044,0.0009938092,0.002173724,0.00023379677],"domain_scores_gemma":[0.98800194,0.0044897674,0.0013357882,0.0033889683,0.00225192,0.00053159933],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025945592,0.0018873267,0.001458622,0.005285238,0.0010453206,0.0021678922,0.0031944287,0.0021717781,0.0053138635],"category_scores_gemma":[0.016786812,0.0013914825,0.0012430488,0.0035507896,0.00091850996,0.004557576,0.003143359,0.0013420543,0.004741416],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001307869,0.0003409022,0.017626768,0.00090106647,0.0004998247,0.0010852246,0.0010045217,0.008869783,0.14900711,0.005617877,0.11284321,0.70089597],"study_design_scores_gemma":[0.0012696794,0.001589284,0.02460304,0.00028688007,0.00055046915,0.0034280247,0.0007296219,0.36980522,0.40383902,0.02601954,0.16723754,0.0006417067],"about_ca_topic_score_codex":0.003419524,"about_ca_topic_score_gemma":0.004589581,"teacher_disagreement_score":0.0053138635,"about_ca_system_score_codex":0.00090687204,"about_ca_system_score_gemma":0.0018730859,"threshold_uncertainty_score":0.017776668},"labels":[],"label_agreement":null},{"id":"W2622881","doi":"","title":"Secrets from the Monster: Extracting Mozilla’s Software Architecture","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":94,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Reverse engineering; Computer science; Software engineering; Software; Architecture; Software system; Source code; Visualization; Software architecture; Programming language; World Wide Web; Artificial intelligence","score_opus":0.012275263809197611,"score_gpt":0.23558381052943744,"score_spread":0.22330854672023984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2622881","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7979425,0.005473101,0.110599026,0.0034917197,0.00049646705,0.0003075346,0.00922011,0.012919846,0.05954973],"genre_scores_gemma":[0.9043632,0.0012594555,0.06474747,0.0001316679,0.00010939026,0.000088497174,0.007975631,0.0005355321,0.020789176],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99943846,0.00007393725,0.0000592656,0.00007000718,0.00029624268,0.0000621382],"domain_scores_gemma":[0.99858356,0.00028594115,0.00021116696,0.00042959943,0.00041150633,0.000078173674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006348072,0.00034158622,0.00020016533,0.004154707,0.0004638293,0.0009964249,0.00025137726,0.00030517197,0.0041556912],"category_scores_gemma":[0.005921187,0.00018968125,0.00023642981,0.002447274,0.00028001334,0.0013989733,0.0012084352,0.00029388486,0.0019071283],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035239613,0.000047836827,0.0861048,0.00028073014,0.00005535041,0.0007167676,0.0018523453,0.0021843924,0.016386596,0.009176174,0.021758966,0.8610836],"study_design_scores_gemma":[0.00006697614,0.00027797185,0.27866685,0.00046337195,0.00017625713,0.0038379235,0.0037856717,0.12241743,0.055784095,0.025265645,0.5091362,0.00012165425],"about_ca_topic_score_codex":0.0067738337,"about_ca_topic_score_gemma":0.008006152,"teacher_disagreement_score":0.0067738337,"about_ca_system_score_codex":0.0005943332,"about_ca_system_score_gemma":0.0007737844,"threshold_uncertainty_score":0.013902187},"labels":[],"label_agreement":null},{"id":"W2623177746","doi":"10.1007/s10664-017-9526-0","title":"An exploratory qualitative and quantitative analysis of emotions in issue report comments of open source systems","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Regione Autonoma della Sardegna","keywords":"Sadness; Gratitude; Recall; Psychology; Emotion classification; Cognitive psychology; Anger; Computer science; Social psychology","score_opus":0.10398874082538329,"score_gpt":0.43043162378672245,"score_spread":0.32644288296133916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2623177746","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9828607,0.000072237395,0.007480012,0.00051895087,0.000083381194,0.00046535846,0.0013239437,0.00007217067,0.007123231],"genre_scores_gemma":[0.98588514,0.000113281334,0.006988532,0.0004754896,0.000099416895,0.0013995364,0.00088921224,0.000097064185,0.004052343],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9893714,0.0066205505,0.00053293695,0.0006436246,0.0021397327,0.0006917541],"domain_scores_gemma":[0.8727664,0.09721487,0.009627407,0.0021049741,0.016275879,0.002010511],"candidate_categories":["scholarly_communication","open_science"],"consensus_categories":[],"category_scores_codex":[0.0078058746,0.00041539533,0.00034136648,0.0024300548,0.0022341902,0.0022539254,0.0006016699,0.0010250864,0.0037153738],"category_scores_gemma":[0.059277672,0.00022871637,0.0002748611,0.0018596709,0.0013544628,0.0016916083,0.0022411519,0.0013919751,0.0008320255],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009806837,0.00079640455,0.09284501,0.0013654702,0.00004147845,0.0010392918,0.8097617,0.00034938805,0.03459059,0.0018293693,0.0049591116,0.051441535],"study_design_scores_gemma":[0.000046003373,0.00091552734,0.2687381,0.0005520976,0.000042908494,0.00048139243,0.696398,0.0017933005,0.009093644,0.0013872547,0.020402689,0.00014903671],"about_ca_topic_score_codex":0.0010704206,"about_ca_topic_score_gemma":0.0023943891,"teacher_disagreement_score":0.99939835,"about_ca_system_score_codex":0.0011743462,"about_ca_system_score_gemma":0.00093372073,"threshold_uncertainty_score":0.04128188},"labels":[],"label_agreement":null},{"id":"W2623394530","doi":"","title":"Improving community awareness in software forges by semantical aggregation of tools feeds","year":2008,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Francophone University Association","funders":"","keywords":"Software; Computer science; Software engineering; Human–computer interaction; Programming language","score_opus":0.024603844632072082,"score_gpt":0.2542580714032772,"score_spread":0.22965422677120514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2623394530","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36303,0.00036000196,0.61912036,0.00087154005,0.000049762177,0.00047272132,0.0017302685,0.008756228,0.0056090113],"genre_scores_gemma":[0.6720332,0.00028916393,0.32159925,0.00006910558,0.000040902218,0.00023133082,0.0036947953,0.0006773811,0.0013647999],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99431133,0.0021487959,0.00045616992,0.0011692576,0.0016547707,0.00025957718],"domain_scores_gemma":[0.97725767,0.011424362,0.0022875594,0.0056037433,0.0024787416,0.0009478637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008690577,0.0007831521,0.00092026143,0.00829656,0.0015896746,0.0042712917,0.0010453978,0.0009943239,0.0008094569],"category_scores_gemma":[0.023065006,0.0009057951,0.0006641564,0.007874136,0.0012321051,0.009655851,0.0058826497,0.0015774468,0.00029612676],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011367437,0.0011160453,0.103370056,0.0012418533,0.00041526678,0.0011269803,0.040645827,0.029518902,0.05692006,0.035304327,0.0066720797,0.72253186],"study_design_scores_gemma":[0.00019688527,0.00056915416,0.15894532,0.0005733665,0.0005758994,0.0012666512,0.02750073,0.43492675,0.115368105,0.16350363,0.096116126,0.00045732642],"about_ca_topic_score_codex":0.006566731,"about_ca_topic_score_gemma":0.007876972,"teacher_disagreement_score":0.008690577,"about_ca_system_score_codex":0.00092974183,"about_ca_system_score_gemma":0.0018553813,"threshold_uncertainty_score":0.045960724},"labels":[],"label_agreement":null},{"id":"W2626335860","doi":"10.6000/1929-7092.2017.06.38","title":"Personal Software Process with Automatic Requirements Traceability to Support Startups","year":2017,"lang":"en","type":"article","venue":"Journal of Reviews on Global Economics","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Traceability; Software engineering; Computer science; Software development; Sequence diagram; Requirements traceability; Software development process; Verification and validation; Software construction; Software; Programming language; Unified Modeling Language; Engineering; Requirement","score_opus":0.05424544222369045,"score_gpt":0.3468016015881474,"score_spread":0.29255615936445695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2626335860","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038866334,0.00044473077,0.9412829,0.0002122474,0.00003531076,0.00054167083,0.00020755465,0.009964495,0.008444616],"genre_scores_gemma":[0.28965265,0.0004231295,0.70366544,0.00009877252,0.000032913013,0.0005208529,0.0009848381,0.0006054949,0.004015927],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99253947,0.0032034244,0.0005942947,0.0009391094,0.0024790352,0.00024468737],"domain_scores_gemma":[0.971951,0.012498082,0.0025439213,0.009537853,0.0031394276,0.0003297224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066260584,0.0007990542,0.000546652,0.0035888797,0.00046134787,0.0014115766,0.0014184976,0.0009622299,0.0041137077],"category_scores_gemma":[0.028544337,0.00051964726,0.00081315154,0.002573426,0.00059999945,0.0029354526,0.0023392364,0.0009350855,0.0017502393],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026166462,0.00059279794,0.0070747454,0.00092513807,0.00010343506,0.00039044657,0.002236201,0.015872777,0.014276938,0.01945184,0.003291377,0.93552274],"study_design_scores_gemma":[0.0003989672,0.0018967214,0.027192526,0.0013195389,0.00042012258,0.0043955026,0.001393008,0.4597372,0.1462446,0.12465327,0.23199947,0.0003490277],"about_ca_topic_score_codex":0.001597377,"about_ca_topic_score_gemma":0.0012973438,"teacher_disagreement_score":0.0066260584,"about_ca_system_score_codex":0.0005436414,"about_ca_system_score_gemma":0.0017661098,"threshold_uncertainty_score":0.035042346},"labels":[],"label_agreement":null},{"id":"W2626932320","doi":"10.1007/s11219-017-9375-5","title":"An investigation of the fault-proneness of clone evolutionary patterns","year":2017,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Cloning (programming); Software maintenance; Similarity (geometry); Biology; Genetics; Computer science; Software system; Software; Gene; Artificial intelligence; Programming language","score_opus":0.049937006655939524,"score_gpt":0.34101401970992457,"score_spread":0.29107701305398503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2626932320","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99495536,0.00007864493,0.004406203,0.00002536024,0.000003267691,0.000013815474,0.0000676342,0.00005374885,0.00039582624],"genre_scores_gemma":[0.9964393,0.000032845164,0.0032199256,0.000004555181,0.0000031013578,0.000007360041,0.00012885622,0.000011104272,0.00015302027],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987037,0.00034975258,0.00013297562,0.00027888775,0.00043365653,0.0001010449],"domain_scores_gemma":[0.93559545,0.046982154,0.008048328,0.0036469887,0.0050249794,0.0007020579],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019203064,0.00021162935,0.0002516408,0.0025000884,0.0003889837,0.0007715722,0.0006280276,0.0006367077,0.0008391184],"category_scores_gemma":[0.035219852,0.00018753065,0.00032202885,0.0016214395,0.0005122606,0.0011305248,0.00040694512,0.00040210361,0.0001104377],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008323896,0.0003883156,0.8001522,0.00025163774,0.0002690602,0.0010898133,0.0018464165,0.028496874,0.042402513,0.0040138685,0.0004365912,0.11982047],"study_design_scores_gemma":[0.00006442275,0.0014837475,0.5945332,0.000068951784,0.00033334622,0.004442239,0.0013903914,0.3620661,0.02623709,0.0081421565,0.0011693354,0.000068995265],"about_ca_topic_score_codex":0.0009755567,"about_ca_topic_score_gemma":0.0011208368,"teacher_disagreement_score":0.0025000884,"about_ca_system_score_codex":0.00033721022,"about_ca_system_score_gemma":0.00033496224,"threshold_uncertainty_score":0.010155678},"labels":[],"label_agreement":null},{"id":"W264881235","doi":"","title":"Avoiding state enumeration in dynamic checking of distributed programs","year":2008,"lang":"en","type":"dissertation","venue":"Spectrum Research Repository (Concordia University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Atomicity; Dependability; Distributed computing; Concurrency; Serialization; Overhead (engineering); Partial order reduction; Code (set theory); Block (permutation group theory); Property (philosophy); Exploit; Programming language; Compiler; Model checking; Set (abstract data type); Database transaction","score_opus":0.021601326375009628,"score_gpt":0.27343948518102557,"score_spread":0.25183815880601595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W264881235","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.061413787,0.00010758216,0.93412524,0.00019876279,0.00002527542,0.00013979255,0.0000928256,0.0021652712,0.0017313805],"genre_scores_gemma":[0.49374115,0.00019424697,0.5030474,0.0001357096,0.000028183807,0.0003633596,0.00036880656,0.00038175794,0.00173931],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9940615,0.0022505661,0.0002990684,0.0009095941,0.0019181246,0.0005610421],"domain_scores_gemma":[0.98193514,0.013057521,0.0012145977,0.0025299208,0.0010396854,0.00022320736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004485022,0.0008161589,0.0010307054,0.0014994186,0.0012273261,0.0025057385,0.00203509,0.00093215937,0.0016858829],"category_scores_gemma":[0.017146902,0.0012068248,0.0018244444,0.0012881556,0.0047946367,0.0042938185,0.0033079754,0.0024622104,0.0003941604],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006460083,0.0002356376,0.006605215,0.00047981195,0.000116072646,0.00047686548,0.0007289181,0.53014785,0.022281548,0.32709917,0.0012650758,0.109917946],"study_design_scores_gemma":[0.00006589165,0.00010926892,0.00026026723,0.000043643442,0.0000474669,0.00006578321,0.000058696813,0.8496842,0.015765343,0.13183703,0.0020306746,0.000031826527],"about_ca_topic_score_codex":0.0056567267,"about_ca_topic_score_gemma":0.007058717,"teacher_disagreement_score":0.0056567267,"about_ca_system_score_codex":0.0020210748,"about_ca_system_score_gemma":0.0036947993,"threshold_uncertainty_score":0.02371937},"labels":[],"label_agreement":null},{"id":"W2669477420","doi":"","title":"Oracle-based Differential Operational Semantics (long version)","year":2016,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure","funders":"","keywords":"Computer science; Programming language; Operational semantics; Semantics (computer science); Oracle; Theoretical computer science; Equivalence (formal languages); Representation (politics); Inference; Artificial intelligence; Mathematics; Discrete mathematics","score_opus":0.013982289425196179,"score_gpt":0.23751358611712536,"score_spread":0.22353129669192917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2669477420","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0057955775,0.0007572429,0.9682585,0.0008886479,0.0003076919,0.00009367245,0.00050233776,0.0028509547,0.02054549],"genre_scores_gemma":[0.5000639,0.0015792483,0.45549545,0.0015416305,0.00057177554,0.00042594288,0.0015517351,0.002206548,0.036563788],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99780756,0.00041842824,0.00024137257,0.00051211775,0.0007741252,0.00024644396],"domain_scores_gemma":[0.9981299,0.0006969099,0.00012961798,0.0005225822,0.0004155823,0.0001054766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026560035,0.00087526167,0.0005323559,0.0011214046,0.00067031477,0.002379152,0.0021062854,0.0012489241,0.011115522],"category_scores_gemma":[0.004890086,0.00050292193,0.00130912,0.0012914785,0.0036775565,0.0055294433,0.0031113517,0.0036676878,0.0025236402],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007595594,0.000029341545,0.00036234103,0.00017654897,0.000018803785,0.0001588849,0.0004713197,0.004081808,0.002585337,0.9441016,0.0048850914,0.0430529],"study_design_scores_gemma":[0.00005036104,0.000095993026,0.0004623309,0.0000891481,0.000053571228,0.00041004372,0.00008865559,0.042010907,0.007519648,0.8342359,0.114919625,0.00006387176],"about_ca_topic_score_codex":0.0028366568,"about_ca_topic_score_gemma":0.0016532581,"teacher_disagreement_score":0.011115522,"about_ca_system_score_codex":0.0019560133,"about_ca_system_score_gemma":0.0010786144,"threshold_uncertainty_score":0.037185133},"labels":[],"label_agreement":null},{"id":"W2677216891","doi":"10.1142/s0218194017500280","title":"Investigating the Effect of Aspect-Oriented Refactoring on the Unit Testing Effort of Classes: An Empirical Evaluation","year":2017,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Unit testing; AspectJ; Computer science; Testability; Regression testing; Programming language; Java; Source code; Software engineering; Software; Aspect-oriented programming; Software development; Reliability engineering; Engineering; Software construction","score_opus":0.05407465647112227,"score_gpt":0.3505202259286955,"score_spread":0.2964455694575732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2677216891","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99255604,0.00035544284,0.0057179527,0.000062537954,0.000011136464,0.00026212487,0.00015293545,0.000083886305,0.000797937],"genre_scores_gemma":[0.9875865,0.00020398438,0.01086372,0.00005366992,0.000023716315,0.00045998432,0.00038426532,0.00006516314,0.00035913082],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9575114,0.02462378,0.003459304,0.0035079194,0.010075501,0.0008220867],"domain_scores_gemma":[0.3583758,0.5628757,0.034586556,0.017553335,0.024403546,0.0022050126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032315742,0.0009251041,0.00067862915,0.0021054829,0.00047378996,0.0009967117,0.0014502875,0.0010525099,0.0010220819],"category_scores_gemma":[0.21723825,0.00042299044,0.0010476031,0.0016860041,0.001431477,0.002010211,0.0011334162,0.0015940754,0.00029423134],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00424146,0.012639595,0.6248129,0.0027164049,0.0013937564,0.000662689,0.008250865,0.030326232,0.0202322,0.0010380454,0.0012156414,0.29247025],"study_design_scores_gemma":[0.00045599873,0.024056762,0.87983674,0.00051810086,0.0009278,0.0006065201,0.0030970057,0.05837179,0.0278519,0.0006546667,0.003467661,0.00015498264],"about_ca_topic_score_codex":0.0014491198,"about_ca_topic_score_gemma":0.0016489376,"teacher_disagreement_score":0.032315742,"about_ca_system_score_codex":0.00091153575,"about_ca_system_score_gemma":0.0010943651,"threshold_uncertainty_score":0.17090404},"labels":[],"label_agreement":null},{"id":"W2727421303","doi":"10.1109/msr.2017.6","title":"Analyzing Program Dependencies in Java EE Applications","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Java; Software engineering; Dependency (UML); Debugging; Container (type theory); Dependency graph; Programming language; Software; Engineering","score_opus":0.029900207443168684,"score_gpt":0.3338960629468567,"score_spread":0.303995855503688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2727421303","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91185033,0.00032228927,0.07566361,0.00017574025,0.000022768963,0.000096872296,0.0009282958,0.00802217,0.0029179873],"genre_scores_gemma":[0.88372505,0.00023902617,0.108237706,0.000101698264,0.000013989919,0.00011631897,0.0037452364,0.0016476157,0.0021733597],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99903,0.00020427589,0.000078106204,0.00022573573,0.00034315707,0.00011872533],"domain_scores_gemma":[0.9933515,0.0046113324,0.0006289374,0.00058551726,0.00071620266,0.00010657551],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007982471,0.000562158,0.00025786797,0.0024467178,0.000671138,0.0005520181,0.0006533204,0.0005868303,0.0007841891],"category_scores_gemma":[0.0061299615,0.0005834892,0.000602095,0.0013928848,0.00047134305,0.001272656,0.000943278,0.0010333197,0.0002282419],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007018719,0.001407643,0.28133246,0.0009616526,0.000286112,0.0028435737,0.0044609415,0.09913237,0.1395718,0.017792005,0.010492706,0.44101694],"study_design_scores_gemma":[0.000047731322,0.00028976763,0.24171947,0.00010387369,0.00021863196,0.0011001255,0.0009719964,0.62818104,0.090781584,0.015382398,0.021075407,0.0001279712],"about_ca_topic_score_codex":0.011028748,"about_ca_topic_score_gemma":0.021106046,"teacher_disagreement_score":0.011028748,"about_ca_system_score_codex":0.0006429876,"about_ca_system_score_gemma":0.00092429575,"threshold_uncertainty_score":0.021929085},"labels":[],"label_agreement":null},{"id":"W2727610827","doi":"10.1007/s10664-017-9529-x","title":"Noise in Mylyn interaction traces and its impact on developers and recommendation systems","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Noise (video); Software; Code (set theory); Source code; Recommender system; Data science; Human–computer interaction; Information retrieval; Software engineering; Data mining; Artificial intelligence; Programming language; Image (mathematics)","score_opus":0.04160546049285644,"score_gpt":0.3472158973043576,"score_spread":0.30561043681150113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2727610827","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9845536,0.00023140195,0.010460142,0.0009901338,0.000064349246,0.000043691707,0.0007374094,0.00076102343,0.0021582274],"genre_scores_gemma":[0.99529195,0.00004520713,0.0021695688,0.00007928289,0.00002677014,0.000032768945,0.0007510285,0.0001372048,0.0014661205],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98703927,0.0054835663,0.00089533365,0.0020456011,0.0038628005,0.00067339465],"domain_scores_gemma":[0.6546255,0.28459486,0.01659839,0.021577993,0.016757784,0.0058454783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009496763,0.0004133096,0.0006758367,0.0027989165,0.0011976777,0.0033463512,0.0012990648,0.0020466072,0.003452421],"category_scores_gemma":[0.24338742,0.00072247244,0.0003283895,0.002222286,0.001235216,0.0044198954,0.0021446378,0.002752805,0.0010701469],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032778112,0.0015106754,0.78810424,0.00029535126,0.00042123647,0.00088814384,0.008076981,0.030168442,0.010289967,0.00994516,0.010560904,0.13646097],"study_design_scores_gemma":[0.00017487626,0.0012080522,0.5566933,0.00021027801,0.00030463206,0.0012503298,0.005659446,0.38907915,0.011429545,0.022721289,0.010994276,0.00027482765],"about_ca_topic_score_codex":0.011901148,"about_ca_topic_score_gemma":0.013385512,"teacher_disagreement_score":0.011901148,"about_ca_system_score_codex":0.0017080551,"about_ca_system_score_gemma":0.0016276629,"threshold_uncertainty_score":0.050224304},"labels":[],"label_agreement":null},{"id":"W2727727095","doi":"10.1109/msr.2017.31","title":"An Empirical Study of the Personnel Overhead of Continuous Integration","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Codebase; Artifact (error); Computer science; Software development; Software; Overhead (engineering); Team software process; Investment (military); Service (business); Order (exchange); Software quality; Software engineering; Software development process; Business; Operating system; Marketing; Finance","score_opus":0.042921716045526706,"score_gpt":0.35156137673846244,"score_spread":0.30863966069293575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2727727095","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9938123,0.00021380349,0.001399022,0.00026962982,0.000016122343,0.000071879025,0.00030908434,0.00004631371,0.0038618473],"genre_scores_gemma":[0.9976993,0.00012763517,0.0009893003,0.000050444858,0.000029044751,0.00006419558,0.00031855903,0.000019890222,0.0007017153],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9855226,0.0071958248,0.0012204392,0.0010541795,0.004108696,0.0008981576],"domain_scores_gemma":[0.585973,0.3082089,0.07035224,0.009747318,0.018558515,0.0071600135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009764764,0.00038937252,0.0002501281,0.0031494624,0.0007504251,0.001472158,0.0012201692,0.000689741,0.003891581],"category_scores_gemma":[0.13195507,0.00038850546,0.00026035393,0.003042645,0.0010918962,0.001741007,0.0010814684,0.0012305205,0.00063765113],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007993866,0.001070769,0.91466904,0.0003320937,0.00013270612,0.00041462158,0.004464782,0.002480044,0.0012042943,0.0008770887,0.0019815133,0.071573704],"study_design_scores_gemma":[0.0000371048,0.0011000506,0.9852844,0.00007563988,0.000055206638,0.0004025359,0.0045998297,0.0044666235,0.00067351444,0.00038281738,0.0028955399,0.00002689315],"about_ca_topic_score_codex":0.0030592605,"about_ca_topic_score_gemma":0.0030619476,"teacher_disagreement_score":0.009764764,"about_ca_system_score_codex":0.0018311598,"about_ca_system_score_gemma":0.0016233169,"threshold_uncertainty_score":0.051641583},"labels":[],"label_agreement":null},{"id":"W2727861051","doi":"10.1109/icpc.2017.28","title":"Bug Report Enrichment with Application of Automated Fixer Recommendation","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Eclipse; Computer science; Ranking (information retrieval); Software bug; Compiler; Metric (unit); Information retrieval; Software; World Wide Web; Software engineering; Programming language; Engineering","score_opus":0.015842011667412743,"score_gpt":0.3068533001828659,"score_spread":0.29101128851545316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2727861051","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31653818,0.005589752,0.55219126,0.00091877603,0.0005654555,0.0011198359,0.0062813116,0.11256434,0.004231199],"genre_scores_gemma":[0.4009989,0.00071928976,0.5810088,0.00027101234,0.00030220294,0.00036978436,0.011231042,0.0009886911,0.0041102753],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960878,0.0009813337,0.00039118962,0.0012541709,0.0011087458,0.0001768237],"domain_scores_gemma":[0.98318577,0.007000956,0.0023512756,0.0021750953,0.004804999,0.0004819163],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002771311,0.0025335841,0.001617864,0.010122248,0.00069212506,0.0010935387,0.001686021,0.0012932485,0.00194166],"category_scores_gemma":[0.022015426,0.00056443986,0.0012102657,0.0036584507,0.00024906243,0.0014910749,0.0008175315,0.0009453957,0.001766448],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006582068,0.0009016681,0.032920513,0.0010310261,0.0006094853,0.0005899619,0.00062116684,0.015863094,0.030327054,0.00038377597,0.021025255,0.89506876],"study_design_scores_gemma":[0.0005049167,0.0013964298,0.046697136,0.00019000673,0.0013722391,0.0012874629,0.0005490841,0.8617769,0.057854183,0.0020883358,0.025972031,0.0003113134],"about_ca_topic_score_codex":0.011518841,"about_ca_topic_score_gemma":0.023434596,"teacher_disagreement_score":0.011518841,"about_ca_system_score_codex":0.0005251597,"about_ca_system_score_gemma":0.0015572811,"threshold_uncertainty_score":0.022903562},"labels":[],"label_agreement":null},{"id":"W2727962479","doi":"10.1109/msr.2017.39","title":"Impact of Continuous Integration on Code Reviews","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code review; Code (set theory); Software; Static program analysis; Quality (philosophy); Software quality","score_opus":0.055824147135579906,"score_gpt":0.3754308338977249,"score_spread":0.319606686762145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2727962479","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97109795,0.0027325714,0.013546548,0.0011769208,0.00022966917,0.00038108608,0.0005761176,0.0018907635,0.008368509],"genre_scores_gemma":[0.99251324,0.00025598137,0.0044810087,0.00014789727,0.00010631226,0.000097341814,0.0005519025,0.00011960261,0.0017266035],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.93309665,0.029443007,0.0037476353,0.008077079,0.023353735,0.0022817585],"domain_scores_gemma":[0.4193955,0.45273376,0.051996645,0.022523396,0.04494725,0.008403437],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026174767,0.001116245,0.001038683,0.003657182,0.0014041789,0.004191236,0.0013180282,0.0014799534,0.0029332945],"category_scores_gemma":[0.26625317,0.00063743937,0.0008940264,0.0021090957,0.0009403797,0.0041137314,0.002481128,0.002120184,0.0013709755],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026382003,0.0017494513,0.5485231,0.0013175078,0.0006094735,0.0012859314,0.003940649,0.018113945,0.013269303,0.0015097135,0.009276643,0.3977661],"study_design_scores_gemma":[0.00020873708,0.0041857744,0.82576305,0.0004345934,0.00070091244,0.0016757891,0.0027299155,0.13586257,0.00962149,0.0030485704,0.015538518,0.0002301251],"about_ca_topic_score_codex":0.005637678,"about_ca_topic_score_gemma":0.005536661,"teacher_disagreement_score":0.9738252,"about_ca_system_score_codex":0.0016846715,"about_ca_system_score_gemma":0.00323487,"threshold_uncertainty_score":0.13842708},"labels":[],"label_agreement":null},{"id":"W2729011926","doi":"10.1109/re.2017.78","title":"Let’s Hear it from RETTA: A Requirements Elicitation Tool for TrAffic Management Systems","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Requirements elicitation; Computer science; Software engineering; Domain (mathematical analysis); Requirements engineering; Requirements analysis; Functional requirement; Domain engineering; Software; Requirements management; Software development; Software requirements; Systems engineering; Software design; Engineering; Software construction","score_opus":0.07787539639908711,"score_gpt":0.33771759420580566,"score_spread":0.2598421978067186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2729011926","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035635144,0.0006126786,0.856114,0.012603972,0.0009672578,0.00091218273,0.00327072,0.06582517,0.02405889],"genre_scores_gemma":[0.17467603,0.00077224465,0.77544326,0.004211187,0.00031620462,0.0009545424,0.0040174443,0.009294391,0.030314686],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99517065,0.00210272,0.0004106118,0.00050272973,0.0015946304,0.00021858023],"domain_scores_gemma":[0.97745585,0.015319353,0.000797544,0.002348723,0.00337799,0.0007005051],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005749696,0.0010794033,0.0008275932,0.0016234276,0.0012424043,0.0032848185,0.0016357539,0.0032808788,0.021083655],"category_scores_gemma":[0.03145578,0.00070507894,0.0011271165,0.0008192601,0.0011097948,0.007013555,0.0026172404,0.0036808655,0.010451767],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009942433,0.0006646943,0.0037446672,0.0025735677,0.0001459225,0.003400242,0.022666616,0.0130860815,0.0951913,0.032795195,0.27786604,0.5468715],"study_design_scores_gemma":[0.00024015998,0.0007021324,0.0029430627,0.0010539402,0.00009778128,0.0027343095,0.0061196387,0.058497384,0.03386405,0.0397358,0.8536058,0.00040593505],"about_ca_topic_score_codex":0.0015874578,"about_ca_topic_score_gemma":0.002561469,"teacher_disagreement_score":0.021083655,"about_ca_system_score_codex":0.000689789,"about_ca_system_score_gemma":0.0012337962,"threshold_uncertainty_score":0.070531845},"labels":[],"label_agreement":null},{"id":"W2729123660","doi":"10.1109/icpc.2017.31","title":"Identifying Code Clones Having High Possibilities of Containing Bugs","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; Software bug; Software maintenance; Computer science; Java; Code (set theory); Cloning (programming); Programming language; clone (Java method); Software evolution; Software; Software system; Software engineering; Biology; Software construction; Genetics","score_opus":0.053378215021452805,"score_gpt":0.32881866081706584,"score_spread":0.27544044579561305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2729123660","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9924102,0.0005471759,0.0057539246,0.000058028258,0.000008909112,0.000066873545,0.0002778244,0.00021927933,0.00065784424],"genre_scores_gemma":[0.990239,0.0002580859,0.008028626,0.000031158295,0.00001355994,0.000055843535,0.0007070958,0.00005419452,0.0006124463],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970778,0.0004627099,0.0003014255,0.0006962918,0.0012130968,0.00024865931],"domain_scores_gemma":[0.9414372,0.026723739,0.016800689,0.004115895,0.009167259,0.0017553304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021107853,0.0005752298,0.0005659922,0.0069119283,0.0006792707,0.0012454274,0.00065575604,0.00092619687,0.0011406586],"category_scores_gemma":[0.031805255,0.00030880584,0.0006693179,0.003068156,0.00072542246,0.0021588074,0.0011779598,0.0004885349,0.00032417924],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023730837,0.00009014755,0.9178007,0.00029306635,0.00011786208,0.001073515,0.0026360657,0.0008272606,0.010579584,0.0006801379,0.00056628615,0.06509797],"study_design_scores_gemma":[0.000020868429,0.00028445243,0.9774828,0.00009473833,0.00021911356,0.0029221298,0.0019893018,0.008536755,0.004812033,0.0012408225,0.0023408819,0.000055980756],"about_ca_topic_score_codex":0.0033899606,"about_ca_topic_score_gemma":0.0041971286,"teacher_disagreement_score":0.0069119283,"about_ca_system_score_codex":0.0006096384,"about_ca_system_score_gemma":0.0006731039,"threshold_uncertainty_score":0.011162996},"labels":[],"label_agreement":null},{"id":"W2730750643","doi":"10.5555/3106039.3106048","title":"A heuristic for estimating the impact of lingering defects: can debt analogy be used as a metric?","year":2017,"lang":"en","type":"article","venue":"Workshop on Emerging Trends in Software Metrics","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); Toronto Metropolitan University","funders":"","keywords":"Technical debt; Computer science; Metric (unit); Heuristic; Software bug; Software; Software quality; Debt; Software development; Business; Artificial intelligence; Finance; Engineering; Operations management","score_opus":0.058286738950395524,"score_gpt":0.3794375124624625,"score_spread":0.32115077351206694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2730750643","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5763453,0.003978445,0.40753475,0.0017389003,0.00012385265,0.0007173074,0.0025123982,0.0021876744,0.004861283],"genre_scores_gemma":[0.85064375,0.00022024555,0.14696984,0.0001343507,0.00006388457,0.0001976316,0.001254054,0.000048679354,0.00046759195],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976484,0.0009516726,0.00027036588,0.0005025816,0.00039787253,0.0002292293],"domain_scores_gemma":[0.97364616,0.020682534,0.0023661675,0.00077096,0.0018986298,0.0006355554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039668437,0.0010753804,0.0013309548,0.006936043,0.000523347,0.0020581589,0.001472334,0.0020782314,0.0009176111],"category_scores_gemma":[0.020226303,0.0004292328,0.0006035848,0.0036420482,0.00085138914,0.0019653123,0.00057525886,0.0009303657,0.0002984628],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074612984,0.0011518195,0.2098758,0.0008888517,0.0005176871,0.0004226795,0.00050215423,0.4398405,0.0041938783,0.0057595964,0.01064217,0.3254587],"study_design_scores_gemma":[0.000078482524,0.0003605572,0.016861476,0.00007268721,0.00008288474,0.00022276396,0.00020204856,0.97354966,0.0013993774,0.005963239,0.001160492,0.000046318473],"about_ca_topic_score_codex":0.0043680314,"about_ca_topic_score_gemma":0.005292395,"teacher_disagreement_score":0.006936043,"about_ca_system_score_codex":0.001881884,"about_ca_system_score_gemma":0.0021692335,"threshold_uncertainty_score":0.020978868},"labels":[],"label_agreement":null},{"id":"W2731817361","doi":"10.4050/f-0073-2017-12033","title":"Benefits and Limitations of Reliance on an Open Architecture Technical Standard to Meet Expectations of an Open System","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"Architecture; Open architecture; Computer science; Open source; Operating system; Software","score_opus":0.07626327144070238,"score_gpt":0.34901665695003736,"score_spread":0.272753385509335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2731817361","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8038121,0.00028965157,0.077329084,0.020493207,0.00019479678,0.00034078068,0.00004532291,0.000350312,0.09714473],"genre_scores_gemma":[0.9791718,0.000109281056,0.018075548,0.0006750198,0.000027305003,0.00011515476,0.000025794157,0.00007587284,0.0017241518],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9156213,0.04288476,0.003762286,0.0018406893,0.03223157,0.0036593685],"domain_scores_gemma":[0.83281004,0.08059807,0.010385394,0.028775325,0.042522505,0.0049087894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08683204,0.00059992855,0.0003419922,0.0012995815,0.0032984414,0.008812937,0.0021388533,0.0020151658,0.0022227617],"category_scores_gemma":[0.12575814,0.00049317325,0.0005584243,0.00088810996,0.0067458875,0.01579438,0.006683946,0.004856394,0.00058280694],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074756733,0.0015955358,0.06934955,0.0009814514,0.00015538417,0.00096806075,0.042780735,0.018217405,0.03511338,0.49564692,0.0068370597,0.32760704],"study_design_scores_gemma":[0.00046906152,0.009572277,0.10308837,0.002715972,0.00036312,0.0022504325,0.19763623,0.07112512,0.054567337,0.4065106,0.15086912,0.00083246006],"about_ca_topic_score_codex":0.003700359,"about_ca_topic_score_gemma":0.00518235,"teacher_disagreement_score":0.08683204,"about_ca_system_score_codex":0.00500331,"about_ca_system_score_gemma":0.010583013,"threshold_uncertainty_score":0.4592172},"labels":[],"label_agreement":null},{"id":"W2732909020","doi":"10.1109/icse-nier.2017.17","title":"Accelerating Software Engineering Research Adoption with Analysis Bots","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Software deployment; Software engineering; Software; Requirements analysis; Software development; Coding (social sciences); Context (archaeology); Data science; World Wide Web; Operating system","score_opus":0.07920962478857292,"score_gpt":0.34161766781605385,"score_spread":0.26240804302748094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2732909020","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21977314,0.0010591089,0.70705146,0.012725843,0.00075713696,0.002838269,0.0002824219,0.03079309,0.02471953],"genre_scores_gemma":[0.47001716,0.00060944434,0.5134134,0.0028732638,0.00021986138,0.0018289948,0.00047191183,0.0025289739,0.008036893],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9558495,0.025366213,0.0020297363,0.004647059,0.010034254,0.0020732284],"domain_scores_gemma":[0.7223136,0.14641693,0.021232273,0.07244328,0.0312447,0.006349204],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.062290207,0.001807173,0.001245138,0.005424659,0.002879487,0.007419174,0.003866229,0.0034517157,0.0044093826],"category_scores_gemma":[0.1779456,0.0019361741,0.0014758412,0.0022136709,0.004772701,0.016247284,0.010727544,0.005556458,0.0031678486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018140378,0.0028249498,0.05608122,0.0015938554,0.00044895077,0.0013671649,0.028568039,0.012557964,0.09145984,0.1180966,0.02982874,0.6553587],"study_design_scores_gemma":[0.0010887203,0.0053810156,0.038050998,0.0030439105,0.0008554254,0.0028598001,0.013960394,0.30402654,0.061224885,0.21902353,0.34956804,0.000916791],"about_ca_topic_score_codex":0.0027172354,"about_ca_topic_score_gemma":0.0033154325,"teacher_disagreement_score":0.9377098,"about_ca_system_score_codex":0.0033594398,"about_ca_system_score_gemma":0.005944055,"threshold_uncertainty_score":0.32942605},"labels":[],"label_agreement":null},{"id":"W2734338314","doi":"10.1109/re.2017.36","title":"What Works Better? A Study of Classifying Requirements","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Non-functional requirement; Requirements engineering; Preprocessor; Machine learning; Artificial intelligence; Usability; Latent Dirichlet allocation; Context (archaeology); Naive Bayes classifier; Non-functional testing; Requirements analysis; Task (project management); Data mining; Requirements management; Support vector machine; Topic model; Software; Software development; Engineering; Programming language; Systems engineering; Human–computer interaction","score_opus":0.11688925302049091,"score_gpt":0.369295255768653,"score_spread":0.2524060027481621,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2734338314","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93791056,0.003538464,0.028674053,0.0063866004,0.00007660618,0.00041139717,0.00089884404,0.00023029649,0.02187331],"genre_scores_gemma":[0.9808663,0.00039975793,0.016700648,0.00043911856,0.000055486995,0.00018688255,0.00055406016,0.00009562566,0.00070217205],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98050356,0.014020699,0.0007481404,0.0019655884,0.0023455657,0.00041644814],"domain_scores_gemma":[0.66469115,0.30730385,0.0125257345,0.005400168,0.008663234,0.0014157913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020278186,0.0005067074,0.0005400129,0.0028821086,0.0011651159,0.0030166726,0.00085852254,0.0009976047,0.0027160882],"category_scores_gemma":[0.14646767,0.00022307283,0.0007291395,0.00385918,0.0017634856,0.007174895,0.00078016816,0.0018303188,0.0005954102],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024081017,0.001805786,0.24633652,0.0022846223,0.00066068565,0.00024505417,0.021598337,0.0066294237,0.0042801383,0.03468703,0.01080438,0.66825986],"study_design_scores_gemma":[0.0007882361,0.004415555,0.59559625,0.0020482447,0.0005961818,0.0014345185,0.03040608,0.1961005,0.008211611,0.11048896,0.0496673,0.00024663805],"about_ca_topic_score_codex":0.0027278622,"about_ca_topic_score_gemma":0.002135518,"teacher_disagreement_score":0.020278186,"about_ca_system_score_codex":0.0021186916,"about_ca_system_score_gemma":0.0009815169,"threshold_uncertainty_score":0.107242584},"labels":[],"label_agreement":null},{"id":"W2734506987","doi":"10.1109/mobilesoft.2017.25","title":"Examining User Complaints of Wearable Apps: A Case Study on Android Wear","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Wearable computer; Android (operating system); Computer science; Wearable technology; Categorization; Android app; Mobile apps; Human–computer interaction; Android application; World Wide Web; Internet privacy; Embedded system; Artificial intelligence","score_opus":0.0894922565214069,"score_gpt":0.3301583812012965,"score_spread":0.24066612467988957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2734506987","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.994635,0.0011984378,0.00206732,0.00038304427,0.000031202264,0.0001306088,0.00037402156,0.000050931172,0.0011294361],"genre_scores_gemma":[0.99121463,0.0013559378,0.0044937884,0.00048998836,0.000079449,0.00016966616,0.0005078486,0.00007199125,0.001616691],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9940772,0.0025261326,0.00065639004,0.00059268036,0.0017620369,0.0003855786],"domain_scores_gemma":[0.94672644,0.036564894,0.0071247886,0.0014797981,0.007146115,0.0009580065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032309084,0.0006476096,0.00069175276,0.0037320976,0.0016852057,0.0013810069,0.0006910207,0.0013296296,0.0006490672],"category_scores_gemma":[0.025169224,0.00035288645,0.00057545956,0.0025688377,0.00091023336,0.0016458365,0.0012054702,0.0009173034,0.0003721455],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005527693,0.0009338738,0.41193923,0.0039962893,0.00023629215,0.041828174,0.36068305,0.00058953284,0.020651372,0.00056389853,0.0110871345,0.14693826],"study_design_scores_gemma":[0.00003710701,0.0012685019,0.6818045,0.0016167599,0.0003213291,0.03841021,0.21680515,0.0053557022,0.009813737,0.0004936601,0.043821115,0.00025218233],"about_ca_topic_score_codex":0.0042796778,"about_ca_topic_score_gemma":0.011912545,"teacher_disagreement_score":0.0042796778,"about_ca_system_score_codex":0.0007814014,"about_ca_system_score_gemma":0.0006383546,"threshold_uncertainty_score":0.017086923},"labels":[],"label_agreement":null},{"id":"W2734655072","doi":"10.4236/jsea.2017.108038","title":"The ISBSG Software Project Repository: An Analysis from Six Sigma Measurement Perspective for Software Defect Estimation","year":2017,"lang":"en","type":"article","venue":"Journal of Software Engineering and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"DMAIC; Benchmarking; Six Sigma; Design for Six Sigma; Software; Computer science; Software project management; Process (computing); Software development; Engineering; Data mining; Systems engineering; Manufacturing engineering; Software construction","score_opus":0.024523588644093604,"score_gpt":0.2930876967016105,"score_spread":0.2685641080575169,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2734655072","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5551607,0.003149051,0.1666786,0.0040390603,0.00033124635,0.0035977513,0.23246787,0.009016472,0.025559224],"genre_scores_gemma":[0.45648116,0.001589298,0.24663065,0.00030329608,0.00012931244,0.0039910334,0.28531444,0.0009021283,0.004658619],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9802917,0.00463876,0.0024873954,0.0016786702,0.010378415,0.0005250271],"domain_scores_gemma":[0.9358803,0.023076879,0.008411158,0.008500548,0.023063134,0.0010678942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013785372,0.00086436153,0.0008741407,0.02571676,0.0006892511,0.002729971,0.0017747866,0.00089632894,0.001809838],"category_scores_gemma":[0.04977086,0.00036240363,0.0011158139,0.034631558,0.0004992004,0.002670287,0.0025156923,0.0012265921,0.0012700586],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038032857,0.00055850536,0.4028462,0.0028785665,0.00047019657,0.000741667,0.0048322286,0.011539725,0.0041469345,0.013258147,0.071496256,0.48685127],"study_design_scores_gemma":[0.00014668626,0.001138328,0.66002536,0.0027052916,0.00044972886,0.0008804248,0.008834127,0.0661886,0.011732736,0.009241674,0.23834454,0.00031241242],"about_ca_topic_score_codex":0.0074509154,"about_ca_topic_score_gemma":0.0071051563,"teacher_disagreement_score":0.02571676,"about_ca_system_score_codex":0.0020188005,"about_ca_system_score_gemma":0.004180445,"threshold_uncertainty_score":0.072904885},"labels":[],"label_agreement":null},{"id":"W2734705500","doi":"","title":"Having Fun With 31.521 Shell Scripts","year":2017,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure","funders":"","keywords":"Scripting language; Computer science; Shell (structure); Programming language; Parsing; Natural language processing; Syntax; POSIX; Artificial intelligence; Engineering","score_opus":0.021416686545107082,"score_gpt":0.24545006909735226,"score_spread":0.22403338255224517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2734705500","genre_codex":"empirical","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45821175,0.002197115,0.22540568,0.0022630845,0.00078833214,0.00063388835,0.1001688,0.12724796,0.08308341],"genre_scores_gemma":[0.55390716,0.00097171677,0.19290885,0.0009126954,0.0001971205,0.00053487485,0.1602327,0.061460067,0.028874919],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99804,0.0004458637,0.00023789715,0.0006534393,0.0004889605,0.00013381614],"domain_scores_gemma":[0.986845,0.0080631,0.00084648776,0.0018970745,0.0020397466,0.0003085192],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019700748,0.00094501744,0.0004759897,0.0022189843,0.0012058428,0.0018330782,0.0006450293,0.000678195,0.0134221725],"category_scores_gemma":[0.018899495,0.000751428,0.0005212511,0.0025226814,0.0011756646,0.0025591552,0.0015703128,0.0012516112,0.010029546],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014769629,0.00030663374,0.048084363,0.002703597,0.00011726431,0.00530711,0.02170948,0.002913672,0.04655886,0.033848092,0.28701323,0.5499608],"study_design_scores_gemma":[0.000060156304,0.00016636813,0.08870383,0.0006824483,0.00008986038,0.006847839,0.0026778725,0.015000459,0.0628956,0.022430368,0.80024403,0.00020108576],"about_ca_topic_score_codex":0.0017457238,"about_ca_topic_score_gemma":0.002556012,"teacher_disagreement_score":0.0134221725,"about_ca_system_score_codex":0.00053399126,"about_ca_system_score_gemma":0.0010236264,"threshold_uncertainty_score":0.04490161},"labels":[],"label_agreement":null},{"id":"W2735411563","doi":"10.1145/3092703.3092722","title":"Lightweight detection of physical unit inconsistencies without program annotations","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of Food and Agriculture; University of Nebraska-Lincoln; University of Toronto; Carnegie Mellon University; U.S. Department of Agriculture; National Science Foundation","keywords":"Computer science; Consistency (knowledge bases); Unit (ring theory); Cyber-physical system; Class (philosophy); Unit testing; Domain (mathematical analysis); Annotation; Static analysis; Artificial intelligence; Programming language; Operating system; Software","score_opus":0.031592640841841754,"score_gpt":0.3212755304476779,"score_spread":0.28968288960583616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2735411563","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.236394,0.000637785,0.70489043,0.00047257752,0.00013183586,0.00032596826,0.0011522743,0.05228696,0.0037081789],"genre_scores_gemma":[0.66695744,0.00018076206,0.32408997,0.00029091572,0.000056278175,0.00026295133,0.002107273,0.004216651,0.0018377115],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.982495,0.003459967,0.0014592602,0.0028901096,0.008893972,0.0008017272],"domain_scores_gemma":[0.91359407,0.036561612,0.0158444,0.023893366,0.009423046,0.00068349333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006712117,0.001517247,0.001313396,0.0058309925,0.0011960235,0.0029747835,0.0027980877,0.001640155,0.0016575088],"category_scores_gemma":[0.06911979,0.0014182288,0.0010792157,0.003480646,0.0016957766,0.005005128,0.0043252013,0.00208683,0.0008328322],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010023837,0.0005529568,0.2871114,0.0018951218,0.0005750394,0.0023843746,0.006778933,0.03758696,0.07523809,0.018573219,0.013077495,0.55522406],"study_design_scores_gemma":[0.00013939949,0.0006858677,0.09485826,0.00064061116,0.00064502074,0.002774513,0.0022167566,0.6273553,0.18099527,0.038655337,0.050613414,0.00042023126],"about_ca_topic_score_codex":0.0034846163,"about_ca_topic_score_gemma":0.005143238,"teacher_disagreement_score":0.006712117,"about_ca_system_score_codex":0.0009917,"about_ca_system_score_gemma":0.0030082217,"threshold_uncertainty_score":0.035497546},"labels":[],"label_agreement":null},{"id":"W273615605","doi":"","title":"Impacts des techniques de construction des ensembles flous sur la précision d'un modèle d'estimation des coûts de logiciels par analogie floue","year":2007,"lang":"fr","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Analogy; Humanities; Mathematics; Artificial intelligence; Philosophy; Computer science; Linguistics","score_opus":0.05441369746616285,"score_gpt":0.31144283179665205,"score_spread":0.2570291343304892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W273615605","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19741733,0.0006851175,0.79581714,0.0002200089,0.000046426587,0.00009155343,0.000113889575,0.00091511663,0.0046935095],"genre_scores_gemma":[0.76434755,0.00037296765,0.23289293,0.000039448707,0.000016420094,0.00011806412,0.00014858467,0.000103718034,0.0019603595],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949425,0.0017470119,0.00028497542,0.0011026149,0.0017407453,0.00018203082],"domain_scores_gemma":[0.986584,0.0086920075,0.0008649587,0.0017079959,0.002019983,0.00013115512],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052068974,0.0010453913,0.0009215974,0.0025353502,0.0008697316,0.0026920254,0.0011368765,0.0010228388,0.0023190859],"category_scores_gemma":[0.023953566,0.0005761066,0.0011085586,0.0015828739,0.000895908,0.0032549084,0.0016322209,0.0011371045,0.00048580623],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005813403,0.00013871328,0.0110812895,0.0003320363,0.0002770712,0.00018084145,0.0013192486,0.438135,0.01409061,0.019471401,0.00062131416,0.5137711],"study_design_scores_gemma":[0.000032442676,0.0003150712,0.009025303,0.00010571333,0.00013462578,0.00019283737,0.00039645203,0.9559445,0.017730027,0.011420716,0.0046352902,0.00006698996],"about_ca_topic_score_codex":0.013528713,"about_ca_topic_score_gemma":0.008073351,"teacher_disagreement_score":0.013528713,"about_ca_system_score_codex":0.0017242087,"about_ca_system_score_gemma":0.001538388,"threshold_uncertainty_score":0.027537048},"labels":[],"label_agreement":null},{"id":"W2738381526","doi":"10.1007/s10664-017-9533-1","title":"EnTagRec ++: An enhanced tag recommendation system for software information sites","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":81,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Ask price; Computer science; Set (abstract data type); Software; Recall; Precision and recall; Information retrieval; World Wide Web; Operating system","score_opus":0.03528399466237,"score_gpt":0.3128422166076763,"score_spread":0.2775582219453063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2738381526","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04843058,0.0016599514,0.29895142,0.00066140044,0.0005055787,0.0012361998,0.057573564,0.57869756,0.0122838095],"genre_scores_gemma":[0.17896825,0.0012516397,0.6316577,0.0008744651,0.00042500222,0.00084158266,0.1246404,0.008985099,0.05235577],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99859816,0.00024722892,0.00011678067,0.00026940042,0.00065715855,0.00011126765],"domain_scores_gemma":[0.9968401,0.00080898387,0.00022914103,0.0010149949,0.0007411945,0.000365675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013482538,0.0021364326,0.0018209654,0.0065635904,0.0008174213,0.0017517938,0.0018668283,0.0014645249,0.018292231],"category_scores_gemma":[0.00472981,0.00075558375,0.0010858891,0.0042024353,0.00015550134,0.0025984652,0.0019322314,0.0010509003,0.024486851],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025710496,0.0007895195,0.009116318,0.00094009127,0.0005154988,0.00049185724,0.00022751784,0.003730399,0.031756803,0.0017389634,0.25961584,0.6885062],"study_design_scores_gemma":[0.0010483483,0.0014836121,0.027508007,0.00019274934,0.00084844657,0.0015026416,0.00048476402,0.5298834,0.09586393,0.009831788,0.3305325,0.0008197309],"about_ca_topic_score_codex":0.0101200165,"about_ca_topic_score_gemma":0.024547275,"teacher_disagreement_score":0.018292231,"about_ca_system_score_codex":0.00055885455,"about_ca_system_score_gemma":0.000764681,"threshold_uncertainty_score":0.061193585},"labels":[],"label_agreement":null},{"id":"W2739015512","doi":"","title":"Combining Qualitative and Quantitative Software Process Evaluation: A Proposed Approach","year":2016,"lang":"en","type":"article","venue":"Roczniki Kolegium Analiz Ekonomicznych","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; Université du Québec à Montréal","funders":"","keywords":"Capability Maturity Model Integration; Computer science; Personal software process; Software engineering; Process (computing); Software Engineering Process Group; Software development process; Software development; Scope (computer science); Software; Verification and validation; Software project management; Process management; Software construction; Engineering; Operations management; Operating system","score_opus":0.06098973076182307,"score_gpt":0.3610039828450392,"score_spread":0.30001425208321614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2739015512","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020372532,0.00022638415,0.99022806,0.00065782375,0.000043760017,0.0010456952,0.00006522375,0.00038026707,0.005315565],"genre_scores_gemma":[0.054326065,0.00020615103,0.9422878,0.00015575692,0.00003974107,0.0018973411,0.00008664664,0.000056664943,0.0009438981],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.91240144,0.055905428,0.004683704,0.0047998554,0.021153431,0.0010561176],"domain_scores_gemma":[0.9199668,0.047107287,0.0046520643,0.0059809904,0.021120666,0.0011722731],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0630618,0.003000271,0.0018948663,0.020067759,0.0022547655,0.010380904,0.0038650064,0.0034217755,0.005414452],"category_scores_gemma":[0.059648227,0.0018544031,0.002415041,0.00958605,0.0062729693,0.009833584,0.008291728,0.0023580976,0.0013386473],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028087987,0.0009018018,0.010183827,0.0034678676,0.0005062107,0.00045547605,0.010551053,0.013733077,0.012487539,0.30441755,0.0031037286,0.63991106],"study_design_scores_gemma":[0.00046394588,0.0015514201,0.0094880015,0.0046279673,0.0008766645,0.001854294,0.018414536,0.4123882,0.021615664,0.4373556,0.09065277,0.0007108879],"about_ca_topic_score_codex":0.002240756,"about_ca_topic_score_gemma":0.0019036952,"teacher_disagreement_score":0.93693817,"about_ca_system_score_codex":0.005551607,"about_ca_system_score_gemma":0.009402214,"threshold_uncertainty_score":0.3335067},"labels":[],"label_agreement":null},{"id":"W2739559489","doi":"10.1155/2017/1046161","title":"A Language and Preprocessor for User-Controlled Generation of Synthetic Programs","year":2017,"lang":"en","type":"article","venue":"Scientific Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; Qualcomm","keywords":"Preprocessor; Computer science; Code (set theory); Programming language; Range (aeronautics); Natural language processing; Code generation; Artificial intelligence; Key (lock); Operating system; Set (abstract data type)","score_opus":0.039218891534402764,"score_gpt":0.3076258757506566,"score_spread":0.2684069842162538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2739559489","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005099923,0.000033078057,0.9484867,0.000080528946,0.000038696766,0.00026769534,0.00088383665,0.043278407,0.0018310911],"genre_scores_gemma":[0.054513466,0.00011099591,0.91998243,0.00023571138,0.00003570448,0.0013055056,0.003021344,0.017580269,0.0032145327],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981013,0.0005909824,0.00023425501,0.00034659432,0.000627026,0.000099822675],"domain_scores_gemma":[0.9909328,0.0056908354,0.00055571785,0.0013823172,0.0012255192,0.00021287682],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003587173,0.0010092873,0.0005313691,0.00089549023,0.0005427679,0.0015740291,0.0021132445,0.00081582763,0.013294138],"category_scores_gemma":[0.014272458,0.001011683,0.0010368174,0.00057934714,0.0010214419,0.0016948155,0.0017664429,0.0021179698,0.0040696175],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016613295,0.00076776574,0.010044263,0.0022364545,0.00029024432,0.0024507684,0.0041193715,0.09210519,0.16410823,0.1231455,0.12543674,0.47363406],"study_design_scores_gemma":[0.0004114099,0.0005215828,0.0021676833,0.0002971447,0.00009341748,0.0021570446,0.00029138432,0.46682677,0.19834667,0.040838886,0.28778023,0.00026784002],"about_ca_topic_score_codex":0.00065536867,"about_ca_topic_score_gemma":0.00075541856,"teacher_disagreement_score":0.013294138,"about_ca_system_score_codex":0.0004968688,"about_ca_system_score_gemma":0.0014738363,"threshold_uncertainty_score":0.04447329},"labels":[],"label_agreement":null},{"id":"W2740279154","doi":"10.1145/3106237.3106267","title":"Why do developers use trivial packages? an empirical case study on npm","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":147,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Reuse; Computer science; Code (set theory); Code reuse; Simple (philosophy); Software engineering; Web application; Programming language; World Wide Web; Software; Engineering","score_opus":0.13768779639791595,"score_gpt":0.397328624915088,"score_spread":0.2596408285171721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2740279154","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9891207,0.00024223622,0.002709034,0.0013382175,0.00001288085,0.00013202574,0.00014756023,0.00006756947,0.0062299133],"genre_scores_gemma":[0.99121195,0.00025134007,0.004795039,0.00049498066,0.000020195943,0.00018292233,0.0002519737,0.00009448562,0.002697078],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9869021,0.0072459485,0.0007591223,0.0011806235,0.0031327312,0.0007795085],"domain_scores_gemma":[0.7883483,0.16445766,0.017003367,0.012147886,0.014261105,0.0037816018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01451749,0.000391702,0.0003847367,0.0023638615,0.0034266799,0.0017991207,0.0018156413,0.0021178538,0.0030505527],"category_scores_gemma":[0.107682906,0.0005208066,0.00032634492,0.002968319,0.0029195896,0.0048009297,0.0028861442,0.00284639,0.00094787544],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006604788,0.0053075184,0.45234564,0.0011455809,0.00008598004,0.015338026,0.31356415,0.001442467,0.0029971157,0.014403146,0.021146974,0.17156292],"study_design_scores_gemma":[0.00035727702,0.0017644184,0.51073116,0.0013258662,0.00016012494,0.011128269,0.31231374,0.01816901,0.00623493,0.01766053,0.1199101,0.00024464328],"about_ca_topic_score_codex":0.008576366,"about_ca_topic_score_gemma":0.014451018,"teacher_disagreement_score":0.01451749,"about_ca_system_score_codex":0.0022592635,"about_ca_system_score_gemma":0.0028133485,"threshold_uncertainty_score":0.07677674},"labels":[],"label_agreement":null},{"id":"W2741187818","doi":"","title":"An Architecture for Effort Estimation of Solutions Based on Open Source","year":2017,"lang":"pl","type":"article","venue":"Roczniki Kolegium Analiz Ekonomicznych","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Open source; Open source software; Architecture; Software; Estimation; Software engineering; Software development; Operating system; Systems engineering; Engineering","score_opus":0.039987341436506556,"score_gpt":0.3298549548518306,"score_spread":0.289867613415324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2741187818","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071328804,0.00017716663,0.98209333,0.00016131927,0.000027400134,0.00025095852,0.00021828976,0.008035855,0.0019027264],"genre_scores_gemma":[0.15589254,0.00025307972,0.83959466,0.000049013575,0.000039415827,0.000501204,0.0012607593,0.00031947394,0.0020897796],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9944265,0.0010844867,0.0007126683,0.0012036317,0.0022301064,0.00034261704],"domain_scores_gemma":[0.99284863,0.0014424921,0.0008624376,0.0015207162,0.0030011537,0.0003245545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005193126,0.0017742866,0.0014041301,0.009542267,0.0012043226,0.0040433067,0.0027309046,0.0012969552,0.003816147],"category_scores_gemma":[0.011547749,0.0009820272,0.0012901224,0.00483974,0.000690526,0.004307221,0.003677635,0.0013666227,0.0020768514],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003977642,0.0005189879,0.023537876,0.0006127568,0.00033209758,0.00031585243,0.0013340112,0.06405253,0.023463978,0.036421306,0.0067768586,0.842236],"study_design_scores_gemma":[0.00006702787,0.00036858086,0.013836448,0.00029841813,0.00025017513,0.00036375388,0.00048309332,0.8907214,0.032221537,0.03748505,0.023713982,0.00019070148],"about_ca_topic_score_codex":0.0055421563,"about_ca_topic_score_gemma":0.004713752,"teacher_disagreement_score":0.009542267,"about_ca_system_score_codex":0.0013572908,"about_ca_system_score_gemma":0.002440564,"threshold_uncertainty_score":0.027464211},"labels":[],"label_agreement":null},{"id":"W2742396485","doi":"10.7287/peerj.preprints.3123v1","title":"Finding and correcting syntax errors using recurrent neural networks","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Syntax error; Syntax; Security token; Abstract syntax; Programming language; Parsing; Artificial intelligence; Natural language processing; Language model; Abstract syntax tree","score_opus":0.06346682147086354,"score_gpt":0.32963931475127095,"score_spread":0.26617249328040743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2742396485","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29611352,0.0008940652,0.6529339,0.0011475108,0.00031999906,0.00011608984,0.0009035751,0.044628505,0.0029428573],"genre_scores_gemma":[0.72846925,0.00029686367,0.26467022,0.00030959008,0.000045413508,0.00008694878,0.0016010203,0.0009202159,0.0036006132],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985544,0.00036987563,0.00010549531,0.0005212114,0.0002992312,0.000149741],"domain_scores_gemma":[0.9948355,0.0024310222,0.0007677586,0.0006244611,0.001202846,0.000138323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014851976,0.0018446487,0.000659371,0.0015349815,0.00044205788,0.0011433101,0.001985542,0.0011749028,0.0015727482],"category_scores_gemma":[0.012448305,0.00077632326,0.00078575657,0.00084407715,0.0005406811,0.0023766705,0.0010845363,0.001630732,0.001112742],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005121682,0.00029523825,0.017196907,0.00040613732,0.00027645178,0.0012212733,0.0009932755,0.18903027,0.05741533,0.0026713884,0.012600245,0.71738124],"study_design_scores_gemma":[0.000019409765,0.000069890506,0.001624426,0.00004146337,0.000068369045,0.000105525294,0.00011702216,0.9775704,0.015607158,0.0031451003,0.001600569,0.000030509185],"about_ca_topic_score_codex":0.012360629,"about_ca_topic_score_gemma":0.019236432,"teacher_disagreement_score":0.012360629,"about_ca_system_score_codex":0.0011517473,"about_ca_system_score_gemma":0.0014804299,"threshold_uncertainty_score":0.02457738},"labels":[],"label_agreement":null},{"id":"W2743326897","doi":"10.1109/qrs.2017.41","title":"Predicting Fault-Prone Classes in Object-Oriented Software: An Adaptation of an Unsupervised Hybrid SOM Algorithm","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Computer science; Machine learning; Data mining; Artificial intelligence; Software fault tolerance; Software quality; Software system; Adaptation (eye); Software metric; Source code; Software; Algorithm; Software development; Programming language","score_opus":0.02652344047027128,"score_gpt":0.285626669439664,"score_spread":0.25910322896939275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2743326897","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.112502486,0.00067473715,0.8792129,0.00027243613,0.00013966073,0.00021783232,0.00049837213,0.0042150426,0.0022665672],"genre_scores_gemma":[0.54415834,0.0003715137,0.44951242,0.000264327,0.0001252509,0.00030884487,0.0016464996,0.00025082295,0.0033618745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99934167,0.0001338904,0.000054790216,0.00020537368,0.00017108288,0.00009319463],"domain_scores_gemma":[0.99852484,0.00056987593,0.000114270675,0.00015579052,0.0005574589,0.0000777542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016569896,0.0011748256,0.0012010725,0.0024623135,0.00053701224,0.0008024051,0.002261596,0.0012802222,0.0010317571],"category_scores_gemma":[0.003010551,0.00046549074,0.0014929617,0.0019622217,0.0004192077,0.0011411111,0.0009003742,0.0010062562,0.0006093714],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033390496,0.0004551446,0.013033385,0.0001361697,0.00039072608,0.00014025686,0.00012852832,0.40715742,0.0037856752,0.0011440038,0.0032647087,0.5700301],"study_design_scores_gemma":[0.000008755514,0.000023699575,0.000703712,0.000007054872,0.000013204506,0.000019002198,0.000014451354,0.99770963,0.00060016534,0.00069351227,0.00020099236,0.0000058473956],"about_ca_topic_score_codex":0.012928173,"about_ca_topic_score_gemma":0.015049462,"teacher_disagreement_score":0.012928173,"about_ca_system_score_codex":0.0006309316,"about_ca_system_score_gemma":0.001010145,"threshold_uncertainty_score":0.025705874},"labels":[],"label_agreement":null},{"id":"W2743912071","doi":"10.1109/qrs.2017.55","title":"Automated Performance Deviation Detection across Software Versions Releases","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Granularity; Computer science; Software; Outlier; Standard deviation; Data mining; Interval (graph theory); Anomaly detection; TRACE (psycholinguistics); Confidence interval; Statistics; Artificial intelligence; Mathematics; Operating system","score_opus":0.02154460755476916,"score_gpt":0.3000235910852413,"score_spread":0.2784789835304721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2743912071","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7636004,0.00032693025,0.21979341,0.00017817607,0.000056591656,0.000109854554,0.0017014446,0.01254675,0.0016864731],"genre_scores_gemma":[0.9622203,0.000052317195,0.035310812,0.000021613634,0.000020489388,0.000050951927,0.0016548352,0.0003020874,0.00036655765],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945629,0.0006297909,0.00040683377,0.0015058041,0.002511266,0.00038349206],"domain_scores_gemma":[0.9738061,0.0106157055,0.0047051883,0.0044588656,0.005808332,0.0006058409],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003283437,0.00095101626,0.0008730827,0.004975216,0.0004642292,0.0017540024,0.0010765904,0.0006762449,0.00074651005],"category_scores_gemma":[0.029966872,0.0003370198,0.0005057173,0.003096226,0.00043717955,0.0015549511,0.0013537864,0.001539394,0.0005492289],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008787141,0.000485941,0.36502635,0.00036938893,0.00030852665,0.0007537459,0.0016754855,0.05191172,0.057319187,0.0022574144,0.0046077413,0.51440567],"study_design_scores_gemma":[0.000042478416,0.00066774135,0.31596854,0.000059480466,0.00010863912,0.0009469584,0.0007890206,0.608393,0.06353823,0.0044418154,0.0048960815,0.00014802846],"about_ca_topic_score_codex":0.003109783,"about_ca_topic_score_gemma":0.0026645977,"teacher_disagreement_score":0.004975216,"about_ca_system_score_codex":0.00054231624,"about_ca_system_score_gemma":0.0007607413,"threshold_uncertainty_score":0.01736468},"labels":[],"label_agreement":null},{"id":"W2745533497","doi":"10.1177/0047281617726278","title":"Rhetorical Genres in Code","year":2017,"lang":"en","type":"article","venue":"Journal of Technical Writing and Communication","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Rhetorical question; Documentation; Code (set theory); Computer science; Technical writing; Software documentation; Internal documentation; Source code; Software; World Wide Web; Software development; Sociology; Linguistics; Programming language; Political science; Software development process; Set (abstract data type); Software construction","score_opus":0.046890451092055765,"score_gpt":0.35200799602619354,"score_spread":0.3051175449341378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2745533497","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.438752,0.01052103,0.057346758,0.013500306,0.0010482665,0.00023667248,0.00045825847,0.00033840621,0.47779834],"genre_scores_gemma":[0.9762114,0.0015021784,0.011231617,0.0005690979,0.00049474544,0.00013585016,0.00019882523,0.00023138904,0.009425002],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.99094605,0.0049151727,0.0004684972,0.0006990567,0.0023484293,0.0006227556],"domain_scores_gemma":[0.9321851,0.049591918,0.0076466277,0.0029696042,0.0054076053,0.002199175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005959032,0.0005438018,0.00046878634,0.017430091,0.0072226604,0.01128826,0.0009731102,0.002125759,0.007381713],"category_scores_gemma":[0.04765741,0.00040713817,0.0005270841,0.008278814,0.014011741,0.012332098,0.0057845227,0.00260844,0.0011984434],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001252676,0.000092021175,0.012765135,0.0005169751,0.00002047238,0.0007105954,0.29172975,0.00043191595,0.0032802385,0.62209576,0.004316839,0.06391505],"study_design_scores_gemma":[0.000069708716,0.00018742138,0.0276075,0.0016841057,0.000055815544,0.002713801,0.19290516,0.004944597,0.003901685,0.45236522,0.31344113,0.00012394573],"about_ca_topic_score_codex":0.002055887,"about_ca_topic_score_gemma":0.0024285826,"teacher_disagreement_score":0.017430091,"about_ca_system_score_codex":0.0038437087,"about_ca_system_score_gemma":0.002061105,"threshold_uncertainty_score":0.031514764},"labels":[],"label_agreement":null},{"id":"W2745662907","doi":"","title":"Requirements Effort Estimation: The State of the Practice","year":2016,"lang":"pl","type":"article","venue":"Roczniki Kolegium Analiz Ekonomicznych","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Computer Research Institute of Montréal","funders":"","keywords":"Estimation; Computer science; Project management; Software project management; Duration (music); Task (project management); Process (computing); Software development; Software; Productivity; Software development process; Process management; Engineering management; Business; Engineering; Systems engineering; Software construction; Economics","score_opus":0.021260365959203775,"score_gpt":0.2949084866774937,"score_spread":0.2736481207182899,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2745662907","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27573097,0.39307407,0.20192844,0.058970243,0.0010217772,0.0002738328,0.0009749717,0.00102484,0.06700091],"genre_scores_gemma":[0.78603137,0.14751944,0.0559927,0.0046513774,0.0012720876,0.00028460144,0.001064754,0.00027082706,0.0029128015],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.94700605,0.019120652,0.0048670946,0.010444247,0.017637176,0.00092474784],"domain_scores_gemma":[0.6253554,0.28387773,0.023601098,0.020891933,0.04381611,0.002457692],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.051699582,0.00074189215,0.0012589542,0.007509543,0.001112627,0.007947568,0.0037779321,0.003192377,0.0020972705],"category_scores_gemma":[0.1468295,0.00082013645,0.00072629127,0.008372576,0.0056102467,0.013474531,0.002981438,0.002661505,0.0011947874],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002198655,0.00015434546,0.028103724,0.0039186985,0.00012072397,0.000061269086,0.0067037474,0.002449829,0.0019559541,0.017516457,0.0044480492,0.9343474],"study_design_scores_gemma":[0.00010235897,0.0019713365,0.20808156,0.03637064,0.0006142645,0.0023797085,0.06960826,0.026251208,0.013341084,0.11989485,0.52065957,0.000725116],"about_ca_topic_score_codex":0.0034626173,"about_ca_topic_score_gemma":0.001996268,"teacher_disagreement_score":0.9483004,"about_ca_system_score_codex":0.003013933,"about_ca_system_score_gemma":0.0033675982,"threshold_uncertainty_score":0.27341676},"labels":[],"label_agreement":null},{"id":"W2747474147","doi":"","title":"A Refactoring Technique for Large Groups of Software Clones","year":2017,"lang":"en","type":"dissertation","venue":"Spectrum Research Repository (Concordia University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Concordia University","keywords":"Code refactoring; clone (Java method); Software maintenance; Maintainability; Software evolution; Cloning (programming); Computer science; Software; Code (set theory); Programming language; Software system; Software engineering; Biology; Genetics; Set (abstract data type); Software construction; Gene","score_opus":0.02937693141409226,"score_gpt":0.3021023258651288,"score_spread":0.27272539445103655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2747474147","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.081982866,0.000881132,0.89452964,0.0003101015,0.00010221196,0.00047025937,0.0003796091,0.019685365,0.0016588077],"genre_scores_gemma":[0.13589002,0.00022995814,0.8576462,0.00014257907,0.00003395718,0.00014431491,0.0011083131,0.0011137386,0.0036909329],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99673176,0.0003941379,0.00028057015,0.0011294022,0.0012942031,0.0001700031],"domain_scores_gemma":[0.99098927,0.0021800285,0.0016209142,0.0027293428,0.002260595,0.00021987899],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020284506,0.0014300551,0.0010515738,0.0038978616,0.0011125327,0.0008064232,0.002161857,0.0015015476,0.0018674815],"category_scores_gemma":[0.0089856535,0.0007361395,0.0019254027,0.0024095995,0.0008517202,0.0018596028,0.0016660161,0.0014946865,0.0012512548],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021310883,0.00019496134,0.010807547,0.0004468011,0.00013516561,0.0010298794,0.001612146,0.008247486,0.12752329,0.0027221404,0.005102615,0.8419649],"study_design_scores_gemma":[0.00027817267,0.0015677565,0.027707024,0.00034085402,0.0007746258,0.006099307,0.001161541,0.51105666,0.34787667,0.014796973,0.08805708,0.00028337198],"about_ca_topic_score_codex":0.0042553707,"about_ca_topic_score_gemma":0.0055775684,"teacher_disagreement_score":0.0042553707,"about_ca_system_score_codex":0.0006119862,"about_ca_system_score_gemma":0.0015401813,"threshold_uncertainty_score":0.010727584},"labels":[],"label_agreement":null},{"id":"W2749195958","doi":"10.7287/peerj.preprints.3186v1","title":"Improved query reformulation for concept location using CodeRank and document structures","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Query expansion; Information retrieval; Sargable; Source code; Query optimization; Web query classification; Web search query; Baseline (sea); Task (project management); Software; Code (set theory); Term (time); Query language; Quality (philosophy); Data mining; Search engine; Programming language","score_opus":0.028961210369557015,"score_gpt":0.3266036725533051,"score_spread":0.29764246218374807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2749195958","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0771125,0.0023379214,0.8951651,0.0010605138,0.00021233367,0.0009226113,0.002034782,0.018623699,0.0025304935],"genre_scores_gemma":[0.245656,0.0007246701,0.7404237,0.00033651572,0.00026665902,0.0004284029,0.007496402,0.00075322046,0.0039144037],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99448425,0.0018267765,0.00054240343,0.00081534736,0.0020294003,0.00030189037],"domain_scores_gemma":[0.9878777,0.006094036,0.00089628523,0.0019440372,0.0029474285,0.00024048559],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029935546,0.001690603,0.0020979103,0.0061803893,0.001031307,0.0021106044,0.001865267,0.0013123145,0.0041344315],"category_scores_gemma":[0.018656218,0.0004965944,0.0013806046,0.0047667865,0.0009334437,0.0054223924,0.0020236503,0.002138248,0.002750001],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008349857,0.000776209,0.003140653,0.001047424,0.00015204313,0.00041880566,0.0013400195,0.03637925,0.07515503,0.015052034,0.03584202,0.8298615],"study_design_scores_gemma":[0.00033781526,0.0007932677,0.002755868,0.00007098227,0.000200274,0.0011373996,0.0011322794,0.8710561,0.074274816,0.015533738,0.03252989,0.00017761711],"about_ca_topic_score_codex":0.01027681,"about_ca_topic_score_gemma":0.009594991,"teacher_disagreement_score":0.01027681,"about_ca_system_score_codex":0.0014713114,"about_ca_system_score_gemma":0.0030200118,"threshold_uncertainty_score":0.020434022},"labels":[],"label_agreement":null},{"id":"W2754109047","doi":"10.1109/tse.2017.2750682","title":"Expanding Queries for Code Search Using Semantically Related API Class-names","year":2017,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Identifier; Programming language; Class (philosophy); Information retrieval; Natural language; Java; Code (set theory); Natural language user interface; World Wide Web; Natural language processing; Artificial intelligence","score_opus":0.039927966108913644,"score_gpt":0.30923921623353773,"score_spread":0.2693112501246241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2754109047","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29748088,0.0054456163,0.5843075,0.0038294205,0.0003425226,0.0028516848,0.018412808,0.07399307,0.013336559],"genre_scores_gemma":[0.2913313,0.0012682987,0.6688223,0.001077013,0.00016160488,0.0007812011,0.030532772,0.0020394365,0.003986091],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99453545,0.0014162061,0.00064384757,0.0012735454,0.0017821575,0.00034871956],"domain_scores_gemma":[0.98820484,0.0074614594,0.00082477124,0.0011521028,0.0019869842,0.00036983305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026271895,0.0025123823,0.002023084,0.009228573,0.0013007214,0.0022660648,0.001906298,0.0022502698,0.0037364725],"category_scores_gemma":[0.018958285,0.0008667045,0.0019961589,0.0049021733,0.0011723455,0.006116499,0.003577225,0.0018436672,0.00258046],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015095903,0.0014732654,0.037043616,0.0045444565,0.00035185475,0.0033135298,0.008927234,0.018697374,0.11824355,0.013744006,0.06959958,0.7225519],"study_design_scores_gemma":[0.00047765987,0.0011398977,0.030684577,0.0007037308,0.00064114784,0.0064420407,0.00820321,0.6208719,0.07731723,0.03078446,0.22229438,0.00043974057],"about_ca_topic_score_codex":0.011053401,"about_ca_topic_score_gemma":0.01491992,"teacher_disagreement_score":0.011053401,"about_ca_system_score_codex":0.001478537,"about_ca_system_score_gemma":0.0029249627,"threshold_uncertainty_score":0.02197814},"labels":[],"label_agreement":null},{"id":"W2754654372","doi":"10.2308/isys-51908","title":"Limiting the Search Space during Controls Evaluation of a Modified Information System","year":2017,"lang":"en","type":"article","venue":"Journal of Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Limiting; Computer science; Audit; Reliability (semiconductor); Space (punctuation); Control (management); Quality (philosophy); Control system; Information space; Quality assurance; Information system; Reliability engineering; Operations management; Artificial intelligence; Engineering; World Wide Web; Accounting; Operating system; Business","score_opus":0.04619519054294601,"score_gpt":0.2987912656962301,"score_spread":0.2525960751532841,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2754654372","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5265845,0.00051534135,0.45168173,0.0009271112,0.00006093521,0.00089870614,0.00016375633,0.0031291707,0.01603868],"genre_scores_gemma":[0.86603814,0.000072296716,0.13143812,0.0001067588,0.000013079508,0.00012760307,0.00015945215,0.0001879621,0.0018565722],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98662853,0.0068884627,0.0006273766,0.0010942358,0.003902524,0.0008589688],"domain_scores_gemma":[0.9514528,0.03431313,0.0034361342,0.0037561534,0.0061538997,0.0008878825],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007590595,0.0009118584,0.0011901922,0.0021295524,0.0011668646,0.0031669333,0.0015237123,0.0012102978,0.0037580477],"category_scores_gemma":[0.06103675,0.000808928,0.00075623643,0.0009837605,0.0015511733,0.0041943756,0.0029838048,0.001040362,0.0005649514],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031164188,0.00071351626,0.021866478,0.00042915947,0.00016740437,0.0007443257,0.0030594165,0.5025299,0.024046084,0.032622587,0.0048536183,0.4058512],"study_design_scores_gemma":[0.000116788935,0.0005711236,0.0023660478,0.000102088416,0.000059762806,0.00011976163,0.00064404286,0.9683548,0.014107803,0.011045924,0.002457341,0.000054429183],"about_ca_topic_score_codex":0.013069161,"about_ca_topic_score_gemma":0.009228887,"teacher_disagreement_score":0.013069161,"about_ca_system_score_codex":0.00287704,"about_ca_system_score_gemma":0.003979022,"threshold_uncertainty_score":0.04014337},"labels":[],"label_agreement":null},{"id":"W2754869403","doi":"10.1109/compsac.2017.237","title":"The Need for Traceability in Heterogeneous Systems: A Systematic Literature Review","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Traceability; Computer science; Requirements traceability; TRACE (psycholinguistics); Abstraction; Software engineering; Semantics (computer science); Key (lock); Focus (optics); Software system; Artifact (error); Risk analysis (engineering); Systems engineering; Data science; Software; Software development; Engineering; Programming language; Artificial intelligence; Computer security","score_opus":0.024128745975928515,"score_gpt":0.3045496449420788,"score_spread":0.2804208989661503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2754869403","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001449111,0.99253476,0.0029163812,0.0017853326,0.00023405059,0.0002983108,0.00016597148,0.000018766817,0.0005972822],"genre_scores_gemma":[0.0127105145,0.9780914,0.007365621,0.00092508266,0.00009805718,0.00046794474,0.00020425064,0.000014636185,0.00012244296],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.97674054,0.009716728,0.0073947017,0.0015205464,0.0042282213,0.00039928113],"domain_scores_gemma":[0.85201603,0.11703194,0.011545901,0.0028996246,0.015484491,0.0010220617],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02864385,0.001481877,0.0044141896,0.02576834,0.0014232126,0.004687939,0.0020794957,0.0030966261,0.0019856733],"category_scores_gemma":[0.11110233,0.0015221954,0.0038061661,0.020865513,0.0022488325,0.00885789,0.0031590536,0.0024851828,0.0003613566],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010520064,0.00008153795,0.002103351,0.610357,0.002063163,0.00046665297,0.0027383855,0.0009143439,0.0007614702,0.006387327,0.005640103,0.36838144],"study_design_scores_gemma":[0.00005055377,0.00013894739,0.0025608253,0.9012718,0.0073436457,0.00066807296,0.002364445,0.00037666416,0.00046365688,0.0048086736,0.07988037,0.00007230852],"about_ca_topic_score_codex":0.009429362,"about_ca_topic_score_gemma":0.025333175,"teacher_disagreement_score":0.02864385,"about_ca_system_score_codex":0.0048383237,"about_ca_system_score_gemma":0.034054767,"threshold_uncertainty_score":0.15148497},"labels":[],"label_agreement":null},{"id":"W2755579216","doi":"10.1007/s10664-017-9547-8","title":"Inference of development activities from interaction with uninstrumented applications","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Documentation; Software engineering; Software development; Generalizability theory; Software; Event (particle physics); Set (abstract data type); Data science; Human–computer interaction; Programming language","score_opus":0.03065809452857621,"score_gpt":0.3057166064988471,"score_spread":0.27505851197027087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2755579216","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8743093,0.00042415998,0.11789103,0.00026864948,0.000029172681,0.00008077405,0.0024312472,0.0013085285,0.003257101],"genre_scores_gemma":[0.9818127,0.00009192041,0.015257395,0.000021725069,0.000017339145,0.000032696593,0.0020427594,0.0000718256,0.0006516234],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9963043,0.001461851,0.0002519086,0.00087004044,0.0008118467,0.00030001398],"domain_scores_gemma":[0.9315521,0.055795178,0.0035078286,0.00633849,0.0021317285,0.0006746345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037011064,0.0006183217,0.0004929721,0.0028491956,0.000320156,0.0014831094,0.0012898602,0.001248255,0.0020744433],"category_scores_gemma":[0.04844766,0.00063462247,0.0009213761,0.0016545422,0.0005290706,0.0020674022,0.0011602322,0.0015004462,0.0010465083],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019778209,0.0012358959,0.6411399,0.00046736863,0.00052576844,0.0005642706,0.0013748643,0.05861074,0.012683031,0.007614177,0.0024782557,0.2713279],"study_design_scores_gemma":[0.0000897128,0.00039801674,0.24910839,0.000107609056,0.0002760933,0.00044583605,0.0005363302,0.7070461,0.018589886,0.020247027,0.0030922254,0.00006271121],"about_ca_topic_score_codex":0.006502566,"about_ca_topic_score_gemma":0.0074893734,"teacher_disagreement_score":0.006502566,"about_ca_system_score_codex":0.00065255025,"about_ca_system_score_gemma":0.0009375759,"threshold_uncertainty_score":0.01957357},"labels":[],"label_agreement":null},{"id":"W2757935660","doi":"10.1109/tse.2017.2755005","title":"Revisiting the Performance Evaluation of Automated Approaches for the Retrieval of Duplicate Issue Reports","year":2017,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Information retrieval; Eclipse; Categorical variable; Software; Notation; Data mining; Machine learning; Programming language","score_opus":0.059809278371222435,"score_gpt":0.29953603310299953,"score_spread":0.2397267547317771,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2757935660","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76075083,0.039152455,0.11251883,0.0024396558,0.0022191436,0.0020370795,0.009971064,0.054041635,0.016869359],"genre_scores_gemma":[0.76335955,0.0033396718,0.19954818,0.0007670198,0.00057055627,0.0005544062,0.026106698,0.0014180174,0.004335886],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9527919,0.015054962,0.005800851,0.006950428,0.01757103,0.0018307389],"domain_scores_gemma":[0.8903194,0.060655896,0.006545838,0.018575722,0.02167656,0.002226623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025950145,0.002935118,0.002799891,0.009422928,0.0017801761,0.0045992853,0.0051743933,0.0027004364,0.0019955833],"category_scores_gemma":[0.10055035,0.00075861765,0.0016422325,0.0066332724,0.0018039845,0.0073112487,0.002460938,0.002390605,0.0025070307],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031868487,0.00274815,0.03501901,0.00544623,0.0016704559,0.0005419428,0.0014648016,0.053426187,0.025611794,0.002721983,0.074530184,0.7936324],"study_design_scores_gemma":[0.0011415489,0.0056285677,0.076579966,0.0008575051,0.0012162229,0.0021394491,0.0025308405,0.7455607,0.07486661,0.0043726373,0.084446736,0.0006591755],"about_ca_topic_score_codex":0.033279806,"about_ca_topic_score_gemma":0.025601402,"teacher_disagreement_score":0.033279806,"about_ca_system_score_codex":0.0032798108,"about_ca_system_score_gemma":0.0051368275,"threshold_uncertainty_score":0.13723916},"labels":[],"label_agreement":null},{"id":"W2758181130","doi":"10.1109/rew.2017.25","title":"Evaluation of Tools for Hairy Requirements and Software Engineering Tasks","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Task (project management); Automatic summarization; Context (archaeology); Precision and recall; Recall; Software engineering; Software; Scale (ratio); Software system; Artificial intelligence; Dependability; Human–computer interaction; Programming language; Systems engineering; Engineering","score_opus":0.132136767363288,"score_gpt":0.3629018789609057,"score_spread":0.2307651115976177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2758181130","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80430514,0.0079838,0.1695807,0.0005857842,0.00034962044,0.001351062,0.0013205415,0.006925067,0.007598358],"genre_scores_gemma":[0.8495459,0.0012194721,0.1440589,0.00017537881,0.00011276848,0.0004807994,0.0029128455,0.00039514442,0.0010988399],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9694192,0.011646305,0.004455249,0.0027001326,0.011034047,0.00074515305],"domain_scores_gemma":[0.8451834,0.11619209,0.013476962,0.007246758,0.0150130745,0.0028878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015250582,0.0016310448,0.0013502056,0.009990514,0.0008775962,0.0034820223,0.0020621556,0.0019224124,0.0012638333],"category_scores_gemma":[0.10710527,0.00038539493,0.0015086598,0.0038338306,0.0010595552,0.0052726436,0.0022944652,0.001203713,0.00058053643],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043552094,0.002045997,0.047853205,0.0059886985,0.0014397164,0.0006104491,0.004793372,0.03883873,0.033637963,0.0037855604,0.008583624,0.8480674],"study_design_scores_gemma":[0.0015049683,0.030530876,0.18124908,0.0029989642,0.0024673592,0.0037172863,0.0077246036,0.57871616,0.1290034,0.013634488,0.04747309,0.00097965],"about_ca_topic_score_codex":0.0017972029,"about_ca_topic_score_gemma":0.0017125868,"teacher_disagreement_score":0.015250582,"about_ca_system_score_codex":0.0014122482,"about_ca_system_score_gemma":0.0012780664,"threshold_uncertainty_score":0.08065373},"labels":[],"label_agreement":null},{"id":"W2758669266","doi":"10.1109/re.2017.62","title":"ECrits — Visualizing Support Ticket Escalation Risk","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ticket; Process (computing); Task (project management); Customer intelligence; Decision support system; IBM; Product (mathematics)","score_opus":0.030212680267136627,"score_gpt":0.33065047974412465,"score_spread":0.300437799476988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2758669266","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47293416,0.0041291527,0.2753265,0.0035149264,0.0007375377,0.0010390419,0.0757054,0.119355366,0.04725796],"genre_scores_gemma":[0.77589256,0.0011756271,0.17571282,0.00027801466,0.00013056656,0.00033498244,0.034231573,0.0034146514,0.0088292025],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983753,0.00028212302,0.00016836826,0.00027156615,0.0007733046,0.00012927082],"domain_scores_gemma":[0.9927545,0.003127978,0.0011315995,0.0010009682,0.0015193342,0.00046552916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023103214,0.0016383787,0.0006263565,0.0073226783,0.00044869163,0.003448209,0.001247981,0.0013128201,0.009674916],"category_scores_gemma":[0.012802336,0.00041264328,0.00085723714,0.0035975813,0.00038474187,0.0028009056,0.0022016126,0.0012335747,0.0020124796],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002262253,0.0010005006,0.13511276,0.002156183,0.0006869576,0.0026288584,0.008614625,0.10039568,0.018477395,0.019471407,0.22381826,0.48537517],"study_design_scores_gemma":[0.00022155422,0.00060628937,0.0787011,0.0005866727,0.00024008351,0.0020129415,0.004406472,0.75484353,0.018455837,0.021306697,0.11829609,0.00032266785],"about_ca_topic_score_codex":0.008650898,"about_ca_topic_score_gemma":0.009859916,"teacher_disagreement_score":0.009674916,"about_ca_system_score_codex":0.00068536465,"about_ca_system_score_gemma":0.0009312534,"threshold_uncertainty_score":0.0323658},"labels":[],"label_agreement":null},{"id":"W2758716390","doi":"10.1109/re.2017.72","title":"Optimized Functionality for Super Mobile Apps","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Alberta Innovates - Technology Futures","keywords":"Computer science; Group cohesiveness; Software; Key (lock); Product (mathematics); Set (abstract data type); App store; Application programming interface; World Wide Web; Software engineering; Computer security; Operating system","score_opus":0.03969598121219515,"score_gpt":0.31620091611060214,"score_spread":0.276504934898407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2758716390","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07898172,0.00021783335,0.9119883,0.00033994162,0.00001899459,0.00023021152,0.00027692565,0.00032876362,0.007617291],"genre_scores_gemma":[0.66444945,0.0002081872,0.33046398,0.00008270846,0.000017461613,0.00034822157,0.0003135477,0.00020312748,0.003913387],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975975,0.0008071723,0.00013125056,0.0005526371,0.0006032046,0.00030825104],"domain_scores_gemma":[0.99645126,0.0021308658,0.00048042453,0.0003667116,0.00039422358,0.00017651966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002478566,0.0014018913,0.0010432429,0.0012210686,0.0005825709,0.0023664944,0.0015744697,0.0012873429,0.004594435],"category_scores_gemma":[0.00784333,0.0010962281,0.0012740816,0.0011221723,0.0010940409,0.004027236,0.0016744552,0.0017471445,0.00042542687],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014419953,0.000119636155,0.0029071479,0.00020700136,0.0000515893,0.00032568528,0.00036674636,0.88084394,0.006967606,0.06034697,0.0009334416,0.046786025],"study_design_scores_gemma":[0.0000075411567,0.00008108621,0.00083008,0.000020593709,0.000021339414,0.00006437919,0.00009092352,0.96987367,0.0010051659,0.026836429,0.001151145,0.000017717504],"about_ca_topic_score_codex":0.0048354226,"about_ca_topic_score_gemma":0.006861763,"teacher_disagreement_score":0.0048354226,"about_ca_system_score_codex":0.0019332875,"about_ca_system_score_gemma":0.0015858738,"threshold_uncertainty_score":0.015369952},"labels":[],"label_agreement":null},{"id":"W2759755869","doi":"10.1109/re.2017.61","title":"What do Support Analysts Know About Their Customers? On the Study and Prediction of Support Ticket Escalations in Large Software Organizations","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ticket; Computer science; IBM; Decision support system; Process (computing); Process management; Knowledge management; Engineering; Computer security; Artificial intelligence","score_opus":0.02174722604130133,"score_gpt":0.2903253830654988,"score_spread":0.2685781570241975,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2759755869","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98061174,0.0005938015,0.012376335,0.002498768,0.000019103636,0.00006147698,0.00055570103,0.00032357787,0.0029595445],"genre_scores_gemma":[0.9931098,0.00023703824,0.005587425,0.00019581363,0.000028160923,0.000020817622,0.00054287683,0.000018573637,0.00025948664],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9937723,0.002430463,0.0004968287,0.0009969964,0.0018749799,0.00042835277],"domain_scores_gemma":[0.9120085,0.066349104,0.009059283,0.003034225,0.0077989916,0.0017498382],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059661428,0.0009480584,0.0005544284,0.0032332605,0.0007644671,0.0032438368,0.0012188874,0.0018045907,0.0015950234],"category_scores_gemma":[0.06246526,0.0004844042,0.0004288079,0.0029210644,0.00060566864,0.0051887087,0.0010543866,0.0015319845,0.00087646843],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090267777,0.0012556636,0.6924594,0.00034298186,0.00015542297,0.00050854473,0.008341549,0.03659501,0.0042589633,0.001476133,0.0064277253,0.24727593],"study_design_scores_gemma":[0.00007534128,0.0012292737,0.355748,0.000510579,0.00014083729,0.0008725889,0.015310843,0.59825444,0.010070536,0.008150175,0.009390487,0.00024687912],"about_ca_topic_score_codex":0.009943307,"about_ca_topic_score_gemma":0.0094807455,"teacher_disagreement_score":0.009943307,"about_ca_system_score_codex":0.0010337973,"about_ca_system_score_gemma":0.0010079368,"threshold_uncertainty_score":0.031552374},"labels":[],"label_agreement":null},{"id":"W2760100496","doi":"10.1109/icsme.2017.82","title":"Is it Safe to Uplift this Patch?: An Empirical Study on Mozilla Firefox","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Crash; Channel (broadcasting); Computer science; Geology; Computer security; Telecommunications","score_opus":0.11698751365113515,"score_gpt":0.41895329812071275,"score_spread":0.3019657844695776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2760100496","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9992287,0.00003091683,0.000074221665,0.000118946606,0.0000019638849,0.000032603155,0.000036867626,0.0000036141837,0.00047214056],"genre_scores_gemma":[0.99819535,0.00009248706,0.0004400784,0.00017744878,0.000009357232,0.00009021348,0.000135184,0.000015999583,0.00084394985],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99617106,0.0017768966,0.00025550686,0.0004736972,0.00078763743,0.000535052],"domain_scores_gemma":[0.8662549,0.08920615,0.027116437,0.0036643802,0.0082604,0.005497682],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01020184,0.00035359227,0.00034973587,0.0015081699,0.0021812671,0.0019446329,0.0014522225,0.0012241732,0.0028631398],"category_scores_gemma":[0.05695958,0.0004987462,0.00034443213,0.0012318175,0.002092222,0.0032674945,0.0010722725,0.0033502437,0.0009055596],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077566324,0.0073223105,0.8742198,0.000262783,0.00007994932,0.0012002639,0.08180891,0.00049789663,0.0018052623,0.000876879,0.0038231423,0.027327195],"study_design_scores_gemma":[0.00007575214,0.001296022,0.9197184,0.00013322088,0.000031310374,0.00027611537,0.07037871,0.0032176673,0.0006030673,0.00025281028,0.003962897,0.000054031465],"about_ca_topic_score_codex":0.01374585,"about_ca_topic_score_gemma":0.019770475,"teacher_disagreement_score":0.01374585,"about_ca_system_score_codex":0.0020551786,"about_ca_system_score_gemma":0.0015478875,"threshold_uncertainty_score":0.05395311},"labels":[],"label_agreement":null},{"id":"W2760127551","doi":"10.1109/tse.2017.2757480","title":"On the Use of Hidden Markov Model to Predict the Time to Fix Bugs","year":2017,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Hidden Markov model; Software bug; Software regression; Context (archaeology); Software; Markov model; Predictive modelling; Data mining; Software engineering; Software development; Markov chain; Data science; Software quality; Machine learning; Artificial intelligence; Programming language","score_opus":0.03385569789076816,"score_gpt":0.2492526278261282,"score_spread":0.21539692993536003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2760127551","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2555553,0.0025396084,0.7322895,0.0016186275,0.00022875937,0.00015906396,0.0011158654,0.0038874794,0.0026058317],"genre_scores_gemma":[0.87041736,0.0013454114,0.1232456,0.00030569948,0.00017899515,0.00010753046,0.0019209925,0.00012192202,0.0023565711],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989114,0.0004644015,0.00008749228,0.0002675962,0.00017059923,0.00009842088],"domain_scores_gemma":[0.9855887,0.012562917,0.0005554451,0.00036521786,0.0007701967,0.00015745785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004174621,0.0013218555,0.0010472619,0.0033396138,0.00071847875,0.0011353791,0.0011604715,0.0016228536,0.0010961357],"category_scores_gemma":[0.0112669235,0.0006017321,0.0015745126,0.0018900811,0.00032058515,0.00205049,0.00057972333,0.0017098605,0.0009012617],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077580113,0.0009304576,0.067306064,0.00020963767,0.0005695799,0.00034329316,0.0004082154,0.54484814,0.0035102847,0.003678328,0.0049602953,0.37245995],"study_design_scores_gemma":[0.0000073361816,0.000043019052,0.002042619,0.000015189691,0.000027561431,0.000036673355,0.000019075702,0.9962619,0.0003662425,0.0010048418,0.0001620954,0.000013475017],"about_ca_topic_score_codex":0.041752994,"about_ca_topic_score_gemma":0.03735791,"teacher_disagreement_score":0.041752994,"about_ca_system_score_codex":0.00083078415,"about_ca_system_score_gemma":0.0014708184,"threshold_uncertainty_score":0.08301991},"labels":[],"label_agreement":null},{"id":"W2760730809","doi":"10.1109/re.2017.64","title":"Panel: Context-Dependent Evaluation of Tools for NL RE Tasks: Recall vs. Precision, and Beyond","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Recall; Context (archaeology); Computer science; Precision and recall; Cognitive psychology; Artificial intelligence; Natural language processing; Psychology","score_opus":0.13364891682149058,"score_gpt":0.3628396132054447,"score_spread":0.22919069638395415,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2760730809","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17843622,0.030814776,0.2906107,0.16315536,0.015086258,0.011090663,0.031084867,0.0076594413,0.27206174],"genre_scores_gemma":[0.616226,0.0068548475,0.19840688,0.028754646,0.005489893,0.0065766284,0.022403926,0.0028495954,0.1124376],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98202986,0.007580176,0.001028215,0.0027769883,0.0059853685,0.00059937825],"domain_scores_gemma":[0.891456,0.057427224,0.0035328967,0.00602889,0.03815194,0.0034029884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03812086,0.00154945,0.00095366663,0.0030259758,0.002667911,0.0049626455,0.0025106238,0.006515563,0.024882311],"category_scores_gemma":[0.08293048,0.0005832732,0.0017412556,0.0023198316,0.0012579388,0.0058715586,0.004149736,0.002590006,0.012101816],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00351898,0.0007371689,0.01795409,0.0018438709,0.0005260248,0.00021807177,0.0010000234,0.004694594,0.023586776,0.007486591,0.6725286,0.26590517],"study_design_scores_gemma":[0.0024801306,0.0067117806,0.14390838,0.007983575,0.0017447934,0.001378637,0.00508395,0.0667453,0.15473512,0.06639355,0.5418344,0.0010003838],"about_ca_topic_score_codex":0.0025919483,"about_ca_topic_score_gemma":0.0054084975,"teacher_disagreement_score":0.03812086,"about_ca_system_score_codex":0.002276034,"about_ca_system_score_gemma":0.0021518336,"threshold_uncertainty_score":0.20160472},"labels":[],"label_agreement":null},{"id":"W2763113360","doi":"10.1145/3127005.3127013","title":"The Characteristics of False-Negatives in File-level Fault Prediction","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; False positives and false negatives; Fault (geology); Fault detection and isolation; Artificial intelligence; False positive paradox; Geology; Seismology","score_opus":0.03869439148207753,"score_gpt":0.2878815819606969,"score_spread":0.24918719047861937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2763113360","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90872383,0.0025740368,0.07873781,0.0013910169,0.00024961712,0.00011547601,0.0019917209,0.0029098422,0.0033066848],"genre_scores_gemma":[0.99108446,0.00015964177,0.006225353,0.00027578257,0.000076303506,0.000038635597,0.001315273,0.00020252932,0.00062206737],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9853689,0.0034268124,0.0013295445,0.003623818,0.0054082447,0.00084270554],"domain_scores_gemma":[0.64880466,0.2889114,0.024744943,0.016677732,0.019134631,0.0017266205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012561654,0.0010366382,0.00109019,0.004204477,0.0008236529,0.0017047095,0.0020592608,0.0018689308,0.0011434931],"category_scores_gemma":[0.14949979,0.00054156,0.0006748263,0.0020315305,0.001316623,0.002857902,0.0010366604,0.0017794075,0.0008516074],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090247375,0.00023965431,0.84162384,0.00034797975,0.00017221876,0.0029319255,0.0014335813,0.023736529,0.0041643213,0.0014665718,0.008438496,0.114542395],"study_design_scores_gemma":[0.00007508025,0.00074095756,0.573152,0.0006211477,0.00037850364,0.0130461585,0.0016838539,0.3451737,0.03424798,0.020164473,0.010501243,0.00021496692],"about_ca_topic_score_codex":0.0036473211,"about_ca_topic_score_gemma":0.00333016,"teacher_disagreement_score":0.012561654,"about_ca_system_score_codex":0.0008194307,"about_ca_system_score_gemma":0.00065192743,"threshold_uncertainty_score":0.06643313},"labels":[],"label_agreement":null},{"id":"W2763684087","doi":"10.1007/s10664-017-9551-z","title":"Analyzing a decade of Linux system calls","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Linux kernel; Computer science; Operating system; System call; sysfs; Configfs; Kernel (algebra); GNU/Linux; Set (abstract data type); Source lines of code; Application programming interface; Software engineering; Programming language; Software","score_opus":0.03008028736372726,"score_gpt":0.30459034001024504,"score_spread":0.2745100526465178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2763684087","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98972416,0.0009902295,0.0014299833,0.00093354774,0.000055550154,0.000009442422,0.0011572784,0.00013416274,0.0055655283],"genre_scores_gemma":[0.9940937,0.0003939316,0.00085829495,0.00015322014,0.00007018516,0.000009039553,0.0021723201,0.000093906754,0.0021553664],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9984029,0.00025915026,0.00010129305,0.00027085075,0.00071132724,0.00025448998],"domain_scores_gemma":[0.97801614,0.011932715,0.0033581026,0.0017640007,0.0038823716,0.0010465841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018535074,0.0002364228,0.00021711507,0.0031940113,0.0009005204,0.0018283633,0.00068070326,0.0009003548,0.001636374],"category_scores_gemma":[0.024443066,0.00033888427,0.00022805047,0.0040638857,0.00083755,0.0020935694,0.0009509166,0.0015880677,0.0005866552],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068983604,0.0005964462,0.7072159,0.0003247074,0.0002202102,0.0011163084,0.012993543,0.01203841,0.007311071,0.024222573,0.034359418,0.19891162],"study_design_scores_gemma":[0.000021805658,0.00023135745,0.8840866,0.00023250873,0.00009443003,0.0005874082,0.009572824,0.027069563,0.003839304,0.0057928017,0.06838535,0.00008605983],"about_ca_topic_score_codex":0.025290895,"about_ca_topic_score_gemma":0.037283313,"teacher_disagreement_score":0.025290895,"about_ca_system_score_codex":0.0016852106,"about_ca_system_score_gemma":0.0012001821,"threshold_uncertainty_score":0.050287366},"labels":[],"label_agreement":null},{"id":"W2764146461","doi":"10.1145/3133909","title":"Understanding the use of lambda expressions in Java","year":2017,"lang":"en","type":"article","venue":"Proceedings of the ACM on Programming Languages","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Computer science; Java; Programming language; Functional programming; Empirical research; Source code; Code (set theory); World Wide Web","score_opus":0.16163113067759652,"score_gpt":0.3304427533680927,"score_spread":0.16881162269049618,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2764146461","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9728158,0.00061323383,0.013721131,0.003165048,0.000018055169,0.000034330656,0.00004088368,0.00007979933,0.009511758],"genre_scores_gemma":[0.9900517,0.0005365952,0.0074757957,0.00031733658,0.000014270037,0.000049723607,0.00003878902,0.00013394066,0.0013819755],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9867798,0.008307465,0.000590681,0.0009966824,0.0025677641,0.00075764366],"domain_scores_gemma":[0.93435514,0.049659684,0.007623679,0.0018507439,0.005223094,0.0012876047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014958063,0.00030527674,0.00023371117,0.002977469,0.0018051493,0.0059783435,0.0012408101,0.0011163093,0.00094364356],"category_scores_gemma":[0.04180615,0.00082627934,0.00032349385,0.001824724,0.0064989747,0.011855628,0.0037725826,0.0017969923,0.00022443179],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000057375677,0.00006718089,0.06603796,0.00023989234,0.000016177135,0.00077220605,0.86820066,0.000336739,0.007481238,0.020818032,0.0007889669,0.035183497],"study_design_scores_gemma":[0.000021195887,0.00016676086,0.21234831,0.0013705638,0.000069017595,0.0021662596,0.648507,0.007345257,0.004680371,0.02841665,0.09471467,0.00019390551],"about_ca_topic_score_codex":0.008047362,"about_ca_topic_score_gemma":0.007891876,"teacher_disagreement_score":0.014958063,"about_ca_system_score_codex":0.0033101658,"about_ca_system_score_gemma":0.0030408595,"threshold_uncertainty_score":0.07910675},"labels":[],"label_agreement":null},{"id":"W2765208928","doi":"10.1002/smr.1920","title":"Finding an effective classification technique to develop a software team composition model","year":2017,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Decision tree; Team composition; Machine learning; Software; Software development; Artificial intelligence; Logistic regression; Set (abstract data type); Data mining; Software engineering; Knowledge management","score_opus":0.02926189430131505,"score_gpt":0.3287277575982327,"score_spread":0.29946586329691766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2765208928","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07406665,0.00009863448,0.9230751,0.00048755278,0.000037994574,0.00034718146,0.00010399464,0.0004633739,0.0013195361],"genre_scores_gemma":[0.47012508,0.00012552386,0.5280992,0.000052835298,0.00003634792,0.00064039486,0.00025510538,0.000033584925,0.00063196477],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99588865,0.0022491321,0.00038072575,0.00040120585,0.0008364851,0.00024377975],"domain_scores_gemma":[0.9811338,0.0143619785,0.001145253,0.00060429465,0.0025170674,0.00023756198],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071515935,0.0009306127,0.0010404374,0.004964214,0.0010416707,0.0021601985,0.001402837,0.0010870592,0.0019536521],"category_scores_gemma":[0.026228057,0.00046184615,0.0013721419,0.0024241912,0.00058163604,0.0025228758,0.0010162339,0.0018849344,0.0004986939],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033644575,0.0009017078,0.050298024,0.00047265855,0.00039105923,0.0003748035,0.0015471989,0.4601073,0.0044918726,0.036638223,0.0041174004,0.44032323],"study_design_scores_gemma":[0.000011593655,0.000054364154,0.0011392201,0.00003449725,0.00002963159,0.000034424454,0.0001575642,0.9925801,0.0006942299,0.0048240838,0.0004294145,0.000010950221],"about_ca_topic_score_codex":0.005333232,"about_ca_topic_score_gemma":0.003583681,"teacher_disagreement_score":0.0071515935,"about_ca_system_score_codex":0.0016127019,"about_ca_system_score_gemma":0.0023411212,"threshold_uncertainty_score":0.03782165},"labels":[],"label_agreement":null},{"id":"W2766228130","doi":"10.1109/scam.2017.25","title":"Does the Choice of Configuration Framework Matter for Developers? Empirical Study on 11 Java Configuration Frameworks","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Java; Popularity; Documentation; Variety (cybernetics); Software engineering; Configuration Management (ITSM); Software maintenance; Software; Software system; Programming language","score_opus":0.04925101324713545,"score_gpt":0.369327377374147,"score_spread":0.32007636412701157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2766228130","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9986449,0.00013492616,0.0003430902,0.000076538665,0.0000028781913,0.000015333053,0.00005875333,0.00001267742,0.0007109588],"genre_scores_gemma":[0.9987268,0.000121847166,0.00061369425,0.000027596758,0.000005576812,0.000022672866,0.00018144287,0.000027265,0.0002730191],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9878442,0.0052400203,0.0011052619,0.00178794,0.0030086886,0.0010139002],"domain_scores_gemma":[0.74429566,0.19251804,0.033198744,0.009051936,0.015491673,0.0054439805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015846439,0.00039821802,0.00043059463,0.0032463216,0.0012555795,0.0024984302,0.001068776,0.001002458,0.0014444944],"category_scores_gemma":[0.117008395,0.0005108928,0.00032818323,0.0037475706,0.0014927186,0.0039419304,0.0013566621,0.001307345,0.00046465112],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034861694,0.00042457177,0.9521084,0.0001583626,0.000060799906,0.00052319054,0.008540591,0.00061572733,0.00073713297,0.00061374664,0.0012283026,0.034640588],"study_design_scores_gemma":[0.000043954966,0.00041910232,0.9681264,0.00012517521,0.000059238075,0.0008984141,0.019382415,0.0038239742,0.0009278965,0.000458876,0.005688993,0.00004568118],"about_ca_topic_score_codex":0.003328024,"about_ca_topic_score_gemma":0.006183641,"teacher_disagreement_score":0.015846439,"about_ca_system_score_codex":0.0012670453,"about_ca_system_score_gemma":0.00086808566,"threshold_uncertainty_score":0.083805025},"labels":[],"label_agreement":null},{"id":"W2766557196","doi":"10.1109/scam.2017.26","title":"On the Relationships Between Stability and Bug-Proneness of Code Clones: An Empirical Study","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Commit; Computer science; Code (set theory); Java; Programming language; Stability (learning theory); Software bug; Software maintenance; Perspective (graphical); Empirical research; Type (biology); Software; Software system; Biology; Artificial intelligence; Database; Mathematics; Machine learning; Statistics","score_opus":0.21208280517832365,"score_gpt":0.3893369986924202,"score_spread":0.17725419351409658,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2766557196","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99872404,0.0001534535,0.00064570084,0.000036304777,0.0000016353384,0.000015718899,0.00011951769,0.000014646817,0.000288995],"genre_scores_gemma":[0.9991042,0.00006458838,0.00047942266,0.000009268846,0.000004940669,0.000013092776,0.00021416054,0.0000074594836,0.00010289195],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99219656,0.0026428043,0.0010579199,0.0013308931,0.0023201297,0.00045174992],"domain_scores_gemma":[0.5936787,0.32890436,0.047229625,0.008332435,0.01909339,0.002761477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009037328,0.00041028863,0.0003925521,0.003653988,0.00058342307,0.0013605516,0.000789761,0.0007860613,0.0014830747],"category_scores_gemma":[0.09871946,0.00034036883,0.0005491328,0.0038486724,0.0013422717,0.0025554078,0.00094746915,0.0013155449,0.00031313516],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007785062,0.00012326964,0.9918921,0.000055005974,0.000076642566,0.00015786182,0.00083857303,0.00060119445,0.00030119251,0.000107646774,0.000105882595,0.0056627872],"study_design_scores_gemma":[0.0000067920105,0.00023403448,0.9903694,0.000021230206,0.00005500351,0.00041882135,0.001439345,0.0065743946,0.00045179934,0.0001319658,0.00028177552,0.000015350019],"about_ca_topic_score_codex":0.002863498,"about_ca_topic_score_gemma":0.0028612674,"teacher_disagreement_score":0.009037328,"about_ca_system_score_codex":0.0006097299,"about_ca_system_score_gemma":0.00060892035,"threshold_uncertainty_score":0.04779452},"labels":[],"label_agreement":null},{"id":"W2766871290","doi":"10.1002/smr.1913","title":"Grouping environmental factors influencing individual decision‐making behavior in software projects: A cluster analysis","year":2017,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Taxonomy (biology); WordNet; Software; Semantic similarity; Cluster analysis; Similarity (geometry); Data science; Management science; Knowledge management; Artificial intelligence; Ecology","score_opus":0.023388160734939274,"score_gpt":0.30276814655737144,"score_spread":0.2793799858224322,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2766871290","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96485347,0.000081653496,0.03265488,0.0001172536,0.000014835147,0.00030848637,0.00048798768,0.00008467081,0.0013966355],"genre_scores_gemma":[0.97954744,0.000050371436,0.019106012,0.0000109811635,0.0000055545906,0.00022851543,0.0006665131,0.000018141242,0.00036640273],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9968311,0.0010890111,0.0003574045,0.0004917069,0.00091507856,0.00031563512],"domain_scores_gemma":[0.98890406,0.005866649,0.0011253175,0.00059861277,0.0029372955,0.00056808227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030904983,0.00073917874,0.0006443408,0.0064429506,0.0014401355,0.0021952393,0.00068950147,0.00059664116,0.0015183953],"category_scores_gemma":[0.013955364,0.00023663459,0.0011766718,0.006635494,0.00084989535,0.0010308592,0.001551522,0.0005142868,0.00026605392],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005520787,0.0005828306,0.82880837,0.0004872907,0.00070343027,0.00043372394,0.011675643,0.018883811,0.005228163,0.0051393895,0.0022164823,0.12528877],"study_design_scores_gemma":[0.00003861844,0.00040017997,0.75783575,0.00021045991,0.00043959974,0.00026587734,0.028848879,0.19344598,0.0040627457,0.010385403,0.003864442,0.00020208738],"about_ca_topic_score_codex":0.01116141,"about_ca_topic_score_gemma":0.008044679,"teacher_disagreement_score":0.01116141,"about_ca_system_score_codex":0.0013589393,"about_ca_system_score_gemma":0.0022987176,"threshold_uncertainty_score":0.022192895},"labels":[],"label_agreement":null},{"id":"W2766877844","doi":"10.1007/s11334-017-0306-1","title":"Predicting different levels of the unit testing effort of classes using source code metrics: a multiple case study on open-source software","year":2017,"lang":"en","type":"article","venue":"Innovations in Systems and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; Université du Québec à Trois-Rivières","funders":"","keywords":"Unit testing; Computer science; Software quality assurance; Software engineering; Software reliability testing; Regression testing; Software metric; Software quality; Non-regression testing; Metric (unit); Quality assurance; Manual testing; Quality (philosophy); Source code; Keyword-driven testing; Software construction; Reliability engineering; Software system; Software; Software development; Programming language; Engineering","score_opus":0.11909794167492378,"score_gpt":0.329284302935347,"score_spread":0.21018636126042323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2766877844","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99775726,0.000020747728,0.0018918896,0.000031076182,0.0000023821244,0.0000129330765,0.00007823035,0.00003499582,0.00017042893],"genre_scores_gemma":[0.9961398,0.000015894608,0.0034114383,0.000005890806,0.0000023605544,0.000011912518,0.00018368033,0.000015329753,0.00021363387],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9979938,0.00075728795,0.000114751856,0.00044849663,0.000493806,0.00019180884],"domain_scores_gemma":[0.96056056,0.031127295,0.002749486,0.0015677921,0.0031943757,0.0008004632],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0031747029,0.0008739375,0.0004945714,0.001994403,0.00035760555,0.0010328405,0.0013131052,0.0014848787,0.0005038477],"category_scores_gemma":[0.022181245,0.00037552844,0.00071704853,0.0016221845,0.00051327574,0.0013253255,0.0005155064,0.00086643914,0.00017091025],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001092658,0.00386454,0.6993956,0.00019735476,0.0003872374,0.0020901777,0.0020583929,0.14972302,0.021553112,0.00086607825,0.0012762776,0.11749548],"study_design_scores_gemma":[0.000059340593,0.0012485421,0.29795647,0.000032508957,0.0001534583,0.00040034758,0.0009675296,0.686885,0.01080159,0.00089790113,0.00053282303,0.000064420485],"about_ca_topic_score_codex":0.01262086,"about_ca_topic_score_gemma":0.018545864,"teacher_disagreement_score":0.9968253,"about_ca_system_score_codex":0.0009848595,"about_ca_system_score_gemma":0.0006399134,"threshold_uncertainty_score":0.025094807},"labels":[],"label_agreement":null},{"id":"W2767247175","doi":"10.1109/icsme.2017.64","title":"Revisiting Turnover-Induced Knowledge Loss in Software Projects","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Long tail; Metric (unit); Software; Tacit knowledge; Knowledge worker; Software engineering; Knowledge management; Engineering; Operations management; Operating system; Mathematics; Statistics; Work (physics)","score_opus":0.04999237012701569,"score_gpt":0.3240932414820764,"score_spread":0.2741008713550607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767247175","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91175044,0.0011709735,0.079198696,0.0016225295,0.000060399256,0.00022210377,0.00052038697,0.00023007601,0.005224347],"genre_scores_gemma":[0.9915072,0.00018834633,0.0074451976,0.000100303,0.00003168587,0.00005754289,0.0002703701,0.000039234517,0.00036015484],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9750905,0.009021686,0.0018748625,0.003349333,0.009116012,0.0015477067],"domain_scores_gemma":[0.687114,0.2134285,0.051496476,0.01860596,0.025273936,0.004081085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024924511,0.0007365205,0.0010789206,0.006768953,0.0015432907,0.0044828015,0.0028539519,0.0018636049,0.0014919066],"category_scores_gemma":[0.2210647,0.00060241326,0.0009358396,0.005270241,0.0041657635,0.010490769,0.0057614786,0.0038966674,0.00027647588],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011451961,0.0012110397,0.7109014,0.0008412662,0.00061289273,0.0022846905,0.009827696,0.06879859,0.002569989,0.031960014,0.003690227,0.166157],"study_design_scores_gemma":[0.00015577841,0.0012150451,0.40101388,0.0005678039,0.0003295402,0.0026231126,0.0076876925,0.45831984,0.0049189553,0.11764049,0.0052472684,0.00028057958],"about_ca_topic_score_codex":0.0068572895,"about_ca_topic_score_gemma":0.0047165756,"teacher_disagreement_score":0.024924511,"about_ca_system_score_codex":0.0037319704,"about_ca_system_score_gemma":0.0020351433,"threshold_uncertainty_score":0.13181502},"labels":[],"label_agreement":null},{"id":"W2767269462","doi":"10.1109/icsme.2017.13","title":"An Exploratory Study of Performance Regression Introducing Code Changes","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Regression testing; Performance metric; Software quality; Benchmark (surveying); Commit; Performance prediction; Reliability engineering; Software performance testing; Quality (philosophy); Software regression; Metric (unit); Software bug; Software; Regression analysis; Performance indicator; Software metric; Code (set theory); Regression; Machine learning; Software development; Simulation; Operating system; Statistics; Database; Engineering; Software construction; Operations management","score_opus":0.05274218923814708,"score_gpt":0.3253429721121572,"score_spread":0.27260078287401013,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767269462","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99597543,0.00011533325,0.0024416423,0.00010285468,0.000009759381,0.0001739052,0.00037801138,0.00018971457,0.00061338814],"genre_scores_gemma":[0.9911958,0.00010323911,0.006864113,0.00008480636,0.000023166247,0.00023022988,0.0009338055,0.00009190538,0.00047285808],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98769206,0.00428067,0.0009691106,0.00188026,0.0045602797,0.0006176278],"domain_scores_gemma":[0.7710912,0.16238838,0.032033898,0.011203175,0.021338124,0.0019452388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008890179,0.0005579141,0.00042063295,0.002987636,0.0006792189,0.0011234863,0.0010208319,0.0006078389,0.00059479673],"category_scores_gemma":[0.078555636,0.00039265087,0.0004627948,0.0021347133,0.001232901,0.0014257928,0.0010130887,0.0014085652,0.0002636716],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009153333,0.0025531403,0.77761126,0.0012658912,0.00030162386,0.0034657265,0.02814059,0.006569528,0.0394266,0.0011418393,0.003537292,0.1350712],"study_design_scores_gemma":[0.000061440114,0.0032819994,0.9464607,0.00018389642,0.00013301005,0.0013825141,0.008313961,0.01713278,0.015397533,0.0006657334,0.0068825223,0.00010385788],"about_ca_topic_score_codex":0.0019974273,"about_ca_topic_score_gemma":0.0032733867,"teacher_disagreement_score":0.008890179,"about_ca_system_score_codex":0.00084453216,"about_ca_system_score_gemma":0.0008203408,"threshold_uncertainty_score":0.047016323},"labels":[],"label_agreement":null},{"id":"W2767331170","doi":"10.1109/ase.2017.8115624","title":"Detecting fragile comments","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"McGill University","keywords":"Identifier; Code refactoring; Computer science; Eclipse; Correctness; Programming language; Precision and recall; Feature (linguistics); Java; Set (abstract data type); Syntax; Data mining; Software; Artificial intelligence","score_opus":0.034871201178286425,"score_gpt":0.311380313370458,"score_spread":0.2765091121921716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767331170","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.430742,0.0050363922,0.43565693,0.0017985678,0.001567952,0.0018665058,0.032044206,0.072973505,0.018313939],"genre_scores_gemma":[0.5179931,0.0014319442,0.4032259,0.0011393776,0.00040910885,0.0009570939,0.043660868,0.008363831,0.022818808],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9866701,0.0018376753,0.0016003546,0.0025481423,0.006806789,0.00053693523],"domain_scores_gemma":[0.8626401,0.061090767,0.026074903,0.014011467,0.034896277,0.0012865984],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006319378,0.0017762606,0.0011090144,0.006716064,0.0012949735,0.0018146702,0.001971189,0.0018886009,0.0037927597],"category_scores_gemma":[0.062470444,0.000593511,0.00083794986,0.0028537347,0.0007202856,0.004178423,0.0029099705,0.0014342088,0.0036957064],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014239292,0.0002629113,0.14587986,0.0057241307,0.00029805166,0.0045124427,0.0058285994,0.004367162,0.07432779,0.009245039,0.08583871,0.66229135],"study_design_scores_gemma":[0.00017534196,0.0007324441,0.11198145,0.0025433374,0.0007058166,0.0089181075,0.0055258633,0.103102796,0.25949392,0.015562025,0.4905571,0.0007018342],"about_ca_topic_score_codex":0.0045804754,"about_ca_topic_score_gemma":0.00735255,"teacher_disagreement_score":0.006716064,"about_ca_system_score_codex":0.0010981106,"about_ca_system_score_gemma":0.0022620147,"threshold_uncertainty_score":0.033420444},"labels":[],"label_agreement":null},{"id":"W2767368277","doi":"10.1109/models.2017.25","title":"Heuristic-Based Recommendation for Metamodel — OCL Coevolution","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Metamodeling; Coevolution; Computer science; Heuristic; Space (punctuation); Theoretical computer science; Artificial intelligence; Programming language","score_opus":0.06288844914027365,"score_gpt":0.3391041632655983,"score_spread":0.27621571412532464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767368277","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.078002624,0.0010879384,0.9106732,0.0007116377,0.000098249606,0.00046232552,0.0004941332,0.0027149406,0.0057548955],"genre_scores_gemma":[0.4352668,0.00029926846,0.5593677,0.00043643912,0.00007979597,0.00033988617,0.0010399653,0.00032442127,0.0028457507],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99678,0.0011391179,0.00020741535,0.000672302,0.0010102782,0.00019094626],"domain_scores_gemma":[0.99162287,0.0042871144,0.0007488245,0.0019864386,0.0010651842,0.0002895337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002997029,0.001498493,0.0018913534,0.0034663428,0.00091287424,0.002161391,0.0034701545,0.0025239896,0.0036491167],"category_scores_gemma":[0.016622156,0.0008935678,0.0015067683,0.0033928698,0.0008665419,0.00307987,0.0016970404,0.001683356,0.0007579217],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004093605,0.001023826,0.01095945,0.000555302,0.00053352886,0.00031633157,0.0006363716,0.39861536,0.0083675925,0.019051326,0.011566161,0.5479653],"study_design_scores_gemma":[0.000047671325,0.00008612567,0.0006085373,0.000026971135,0.000060665727,0.00008930979,0.00007314983,0.98805004,0.001490667,0.006646219,0.0027915512,0.000029091445],"about_ca_topic_score_codex":0.01090834,"about_ca_topic_score_gemma":0.02189822,"teacher_disagreement_score":0.01090834,"about_ca_system_score_codex":0.0018546609,"about_ca_system_score_gemma":0.0019226283,"threshold_uncertainty_score":0.021689713},"labels":[],"label_agreement":null},{"id":"W2767380523","doi":"10.1109/icsme.2017.24","title":"Understanding Stack Overflow Code Fragments","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Fragment (logic); Code (set theory); Source code; Domain (mathematical analysis); Code reuse; Code review; Reuse; Programming language; Static program analysis; Software; Engineering; Software development","score_opus":0.16586794679213024,"score_gpt":0.33182580013971075,"score_spread":0.1659578533475805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767380523","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98534644,0.00037989262,0.009273503,0.0012715255,0.000008594691,0.000034283425,0.00007032865,0.00012613424,0.0034893553],"genre_scores_gemma":[0.99310243,0.000323975,0.0051225056,0.0001335006,0.000012641662,0.000019563064,0.00016026045,0.000045435714,0.0010796505],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9964463,0.0014688249,0.00025333255,0.00028945672,0.0013099667,0.00023213404],"domain_scores_gemma":[0.910528,0.06715563,0.011532346,0.0026871064,0.0068098158,0.0012871784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073285485,0.00041884696,0.00027301692,0.0032491123,0.0009928037,0.00250196,0.0007649132,0.0013231882,0.0019307679],"category_scores_gemma":[0.066968724,0.00040933385,0.00029574445,0.0009972408,0.0015542546,0.008605986,0.0025525934,0.0010129112,0.0002700002],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019169545,0.00016149088,0.28106537,0.00044108432,0.000038312006,0.002013909,0.55519176,0.0011252359,0.00927975,0.0055332156,0.0022881443,0.14266995],"study_design_scores_gemma":[0.000042796913,0.0004934426,0.53418356,0.001335451,0.00013698812,0.0041420977,0.35094222,0.025555156,0.007784404,0.021111766,0.054083582,0.00018857687],"about_ca_topic_score_codex":0.0075599416,"about_ca_topic_score_gemma":0.008445132,"teacher_disagreement_score":0.0075599416,"about_ca_system_score_codex":0.0016766024,"about_ca_system_score_gemma":0.0015118172,"threshold_uncertainty_score":0.038757563},"labels":[],"label_agreement":null},{"id":"W2767550398","doi":"10.1109/ase.2017.8115667","title":"Detecting unknown inconsistencies in web applications","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Intel Corporation","keywords":"Computer science; JavaScript; Web application; Programming language; Code (set theory); Source code; Novelty; Matching (statistics); Set (abstract data type); Data mining; Information retrieval; World Wide Web","score_opus":0.02861547933864964,"score_gpt":0.29362654080750267,"score_spread":0.265011061468853,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767550398","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7925321,0.0015781377,0.1927415,0.0004258325,0.000110926456,0.00023362864,0.0006125684,0.010080598,0.0016847133],"genre_scores_gemma":[0.83984166,0.00031392896,0.1574494,0.00016523042,0.00004305306,0.00007969224,0.0009897067,0.0005293383,0.0005881242],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98698145,0.0022484274,0.0013788716,0.0022690007,0.0064757015,0.0006465609],"domain_scores_gemma":[0.938602,0.03430764,0.011333788,0.00566255,0.0094096875,0.0006843112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054735118,0.00090317737,0.0010540988,0.005991043,0.0008572646,0.002064081,0.002399157,0.0015309557,0.00048091708],"category_scores_gemma":[0.046072103,0.0008027958,0.0007968886,0.003028971,0.00085429515,0.0029752022,0.0021894064,0.0012138787,0.000245538],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009187648,0.0006327225,0.42821676,0.0014259819,0.00057041366,0.007936069,0.003944138,0.0292497,0.056027036,0.0060479715,0.004410399,0.46062005],"study_design_scores_gemma":[0.00015646694,0.0006043754,0.11299047,0.0004131285,0.0008158994,0.008298286,0.0021104089,0.7313638,0.109604806,0.018406734,0.015011381,0.00022424014],"about_ca_topic_score_codex":0.0023602883,"about_ca_topic_score_gemma":0.0032895757,"teacher_disagreement_score":0.005991043,"about_ca_system_score_codex":0.0007313755,"about_ca_system_score_gemma":0.0015180664,"threshold_uncertainty_score":0.028947055},"labels":[],"label_agreement":null},{"id":"W2767563557","doi":"10.1109/icsme.2017.17","title":"On-demand Developer Documentation","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":105,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Victoria; University of British Columbia; McGill University","funders":"","keywords":"Documentation; Computer science; Software documentation; Internal documentation; Technical documentation; Face (sociological concept); Knowledge management; World Wide Web; Software; Software engineering; Engineering management; Process management; Software development; Business; Engineering; Software development process; Software construction","score_opus":0.02363040197798909,"score_gpt":0.3161471257476193,"score_spread":0.29251672376963017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767563557","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04170443,0.0030602727,0.70519716,0.027696528,0.0013047745,0.0010134196,0.00086456747,0.0154882725,0.20367059],"genre_scores_gemma":[0.2842435,0.0034180186,0.6016399,0.0066893734,0.001297961,0.0008848577,0.0029791158,0.00569808,0.093149066],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96535957,0.0162347,0.0017473191,0.0032485893,0.011837174,0.0015725236],"domain_scores_gemma":[0.7705458,0.07872922,0.011192603,0.10027879,0.03183634,0.0074172723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023814622,0.0012577142,0.0010880346,0.0040466823,0.0025070284,0.008966769,0.0072211595,0.004534207,0.033670615],"category_scores_gemma":[0.09842211,0.0013915999,0.0009653722,0.0025203738,0.0019450446,0.01496876,0.011219047,0.0039667212,0.022758488],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003007287,0.0010727451,0.008977664,0.0016806668,0.00007165639,0.00096678716,0.0066672764,0.0023212961,0.010867278,0.13756327,0.10314625,0.7263643],"study_design_scores_gemma":[0.00022280622,0.00042234102,0.006870075,0.0021276716,0.000072852636,0.0032585787,0.003692841,0.019789232,0.0128862215,0.13915674,0.8113137,0.00018700605],"about_ca_topic_score_codex":0.0015995031,"about_ca_topic_score_gemma":0.0037216397,"teacher_disagreement_score":0.033670615,"about_ca_system_score_codex":0.002059742,"about_ca_system_score_gemma":0.009461357,"threshold_uncertainty_score":0.12594527},"labels":[],"label_agreement":null},{"id":"W2767608636","doi":"10.1109/softwaremining.2017.8100848","title":"Mining specifications using nested words","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Generality; Scalability; Nested set model; Set (abstract data type); Theoretical computer science; Word (group theory); Constant (computer programming); Automaton; Class (philosophy); Programming language; Data mining; Artificial intelligence; Database; Mathematics","score_opus":0.1601463252538005,"score_gpt":0.34550422555665466,"score_spread":0.18535790030285418,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767608636","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0696921,0.00043741005,0.9196134,0.00022358123,0.000027988477,0.00036248754,0.0028332795,0.005284232,0.0015254604],"genre_scores_gemma":[0.3424062,0.0003990031,0.6407747,0.000132454,0.00003006953,0.0005122685,0.013484877,0.00061099214,0.0016494367],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99598145,0.00072243594,0.00067318784,0.0010071519,0.0013972435,0.00021856753],"domain_scores_gemma":[0.98738575,0.006729318,0.0012795199,0.0020012779,0.00236118,0.00024288375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015879681,0.0012493881,0.0010751765,0.0053980527,0.0008094014,0.0022296314,0.0021211447,0.0013127244,0.0020531456],"category_scores_gemma":[0.019437252,0.00083654135,0.0021875945,0.0032946717,0.0008749799,0.0042056153,0.0020685901,0.0008813303,0.0010541872],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004719658,0.0003811455,0.04203322,0.0016436949,0.0003704353,0.0019928145,0.0027095724,0.15293582,0.031539805,0.05591619,0.0071865264,0.70281875],"study_design_scores_gemma":[0.000056976893,0.00018730029,0.0028899116,0.00015027027,0.00012754864,0.00077310053,0.001041544,0.8466879,0.022849008,0.11039097,0.014769475,0.00007589768],"about_ca_topic_score_codex":0.0067064925,"about_ca_topic_score_gemma":0.009573803,"teacher_disagreement_score":0.0067064925,"about_ca_system_score_codex":0.000998166,"about_ca_system_score_gemma":0.0027570704,"threshold_uncertainty_score":0.01333493},"labels":[],"label_agreement":null},{"id":"W2767683547","doi":"10.1109/icsme.2017.33","title":"Bug Propagation through Code Cloning: An Empirical Study","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; Cloning (programming); Computer science; clone (Java method); Programming language; Code (set theory); Software maintenance; Source code; Java; Software evolution; Software bug; Codebase; Commit; Software; Software system; Biology; Database; Genetics; Software construction","score_opus":0.09509380129141144,"score_gpt":0.4050089130761953,"score_spread":0.30991511178478387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767683547","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9976107,0.00021102221,0.0015133463,0.00005852676,0.0000032860123,0.000058523245,0.00012714347,0.000030423891,0.00038714719],"genre_scores_gemma":[0.99738723,0.00018327862,0.0017958336,0.00003857015,0.0000085988295,0.00005415665,0.0003105798,0.000020855708,0.00020089367],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9878505,0.0051906197,0.0012844353,0.0016824738,0.003502706,0.0004893272],"domain_scores_gemma":[0.68748987,0.23189652,0.044446602,0.0134096835,0.020525662,0.0022317185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01003719,0.00052829477,0.00041412524,0.004355638,0.0010894394,0.001495624,0.0014043334,0.0009221823,0.0008761433],"category_scores_gemma":[0.10988945,0.0004865885,0.0004603307,0.004161404,0.0018857429,0.002747183,0.001598496,0.0016340098,0.00023770053],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012398872,0.00048754457,0.96267563,0.00020797575,0.00010376504,0.0005685397,0.006820359,0.0007258877,0.00089978485,0.00038583588,0.0005615111,0.026439201],"study_design_scores_gemma":[0.000047414745,0.0007920573,0.9653616,0.00019594045,0.00018026865,0.0030368392,0.008062912,0.016176015,0.0022301292,0.00083560566,0.0030172465,0.000064025444],"about_ca_topic_score_codex":0.004657718,"about_ca_topic_score_gemma":0.0060981656,"teacher_disagreement_score":0.01003719,"about_ca_system_score_codex":0.0009457875,"about_ca_system_score_gemma":0.0010716864,"threshold_uncertainty_score":0.053082347},"labels":[],"label_agreement":null},{"id":"W2767782162","doi":"10.1109/ase.2017.8115681","title":"AnswerBot: Automated generation of answer summary to developers' technical questions","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":122,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Automatic summarization; Paragraph; Information retrieval; Question answering; Task (project management); Domain (mathematical analysis); Selection (genetic algorithm); World Wide Web; Java; Key (lock); Relevance (law); Data science; Artificial intelligence","score_opus":0.0399541912071572,"score_gpt":0.32437392081935573,"score_spread":0.2844197296121985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767782162","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08129517,0.002299966,0.6612524,0.0017064121,0.0006798781,0.0030450963,0.01761501,0.22528411,0.0068219644],"genre_scores_gemma":[0.147801,0.0006838912,0.783503,0.0007006252,0.00031007946,0.0015855118,0.050389938,0.0035462962,0.011479672],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978416,0.0008099711,0.00018875554,0.00049951114,0.0005526213,0.00010743066],"domain_scores_gemma":[0.9906949,0.005097882,0.00090232387,0.00084295217,0.0020513325,0.0004106507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030095652,0.00242564,0.0013820567,0.004188206,0.0007408797,0.0016254351,0.0016475901,0.0017634907,0.011175586],"category_scores_gemma":[0.01629208,0.00052415597,0.00088132644,0.0013650216,0.00038861678,0.0026329057,0.0023652248,0.001143077,0.006336438],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011211481,0.00094910985,0.004193042,0.0031286739,0.00020966998,0.00080685917,0.0026775047,0.0052207243,0.062897146,0.0034078169,0.16108277,0.7543055],"study_design_scores_gemma":[0.001062361,0.0033286652,0.017265674,0.0005720191,0.00057679974,0.0014733158,0.0039723115,0.5483208,0.1379853,0.016147634,0.26891902,0.0003761732],"about_ca_topic_score_codex":0.0019892864,"about_ca_topic_score_gemma":0.0032707525,"teacher_disagreement_score":0.011175586,"about_ca_system_score_codex":0.00063290494,"about_ca_system_score_gemma":0.0012462889,"threshold_uncertainty_score":0.037386},"labels":[],"label_agreement":null},{"id":"W2767795225","doi":"10.1007/s10664-017-9559-4","title":"Are tweets useful in the bug fixing process? An empirical study on Firefox and Chrome","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; World Wide Web; Social media; Process (computing); Empirical research; Security bug; Microblogging; Software; Internet privacy; Computer security","score_opus":0.06260857921942889,"score_gpt":0.36454485204685283,"score_spread":0.30193627282742397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767795225","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9988537,0.000071142014,0.00009042032,0.00018729985,0.000008509571,0.000012068556,0.000100223995,0.000008767184,0.0006680115],"genre_scores_gemma":[0.998657,0.00007700311,0.00022405262,0.000077904566,0.000019251733,0.0000140501625,0.0002751934,0.000016602839,0.0006388741],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9979722,0.0009561869,0.00010845163,0.00021812164,0.00047124608,0.00027371632],"domain_scores_gemma":[0.9090018,0.069998525,0.0123163285,0.0018066766,0.004293144,0.0025833913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027399543,0.00032918094,0.00030168326,0.0016547117,0.0012975056,0.00209292,0.00064162485,0.001469494,0.0019238209],"category_scores_gemma":[0.044225994,0.00029885207,0.0002639186,0.001445395,0.00080875953,0.0031976735,0.0007216116,0.001929505,0.0005754585],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008794581,0.001815744,0.9448568,0.0001977052,0.00010171432,0.0006263214,0.01998851,0.00034958214,0.0027441047,0.00076916436,0.0023117836,0.025359066],"study_design_scores_gemma":[0.000054295404,0.0005241237,0.9678115,0.00009776734,0.0001253502,0.00027295083,0.022479633,0.0029883287,0.0015333605,0.00028659913,0.003780279,0.000045687357],"about_ca_topic_score_codex":0.016365439,"about_ca_topic_score_gemma":0.016914064,"teacher_disagreement_score":0.016365439,"about_ca_system_score_codex":0.0010742617,"about_ca_system_score_gemma":0.0009686766,"threshold_uncertainty_score":0.03254038},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"agree"},{"id":"W2768032356","doi":"10.1109/icsme.2017.34","title":"Forecasting the Duration of Incremental Build Jobs","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Duration (music); Computer science; Transparency (behavior); Process (computing); Plan (archaeology); Dependency graph; Software; Point (geometry); Software engineering; Graph; Industrial engineering; Operating system; Engineering; Theoretical computer science","score_opus":0.051099516449853936,"score_gpt":0.2932236149235693,"score_spread":0.24212409847371535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2768032356","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92934704,0.0005987084,0.056499865,0.0002605297,0.00012001516,0.000095611176,0.006192441,0.0045308303,0.002355011],"genre_scores_gemma":[0.9682276,0.00019591494,0.022568101,0.000031360025,0.00002722297,0.000054748223,0.008256611,0.00021281197,0.00042555167],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9989467,0.00012117088,0.00007947482,0.00028445185,0.00045423725,0.000114003546],"domain_scores_gemma":[0.98682225,0.0066633145,0.0018117983,0.0012410949,0.002798175,0.0006634735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018548929,0.00084238505,0.0003997215,0.002853975,0.00030424295,0.00077334733,0.0007923458,0.0005353586,0.00068473135],"category_scores_gemma":[0.016366668,0.0005276788,0.00043305857,0.0016219498,0.00025927118,0.0012588497,0.00054861413,0.0008904711,0.00044208847],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001081203,0.00020063277,0.4334043,0.00044983093,0.00022052492,0.00042981515,0.00096606323,0.37364152,0.01765617,0.002187074,0.010446426,0.15931652],"study_design_scores_gemma":[0.000038846454,0.00022551588,0.13917732,0.000057238907,0.000093348084,0.00016255573,0.00038677646,0.8407585,0.012071578,0.0021237114,0.0048323185,0.00007230875],"about_ca_topic_score_codex":0.009580379,"about_ca_topic_score_gemma":0.013478836,"teacher_disagreement_score":0.009580379,"about_ca_system_score_codex":0.00082556746,"about_ca_system_score_gemma":0.0008000165,"threshold_uncertainty_score":0.019049227},"labels":[],"label_agreement":null},{"id":"W2768062948","doi":"10.1109/ase.2017.8115713","title":"FEMIR: A tool for recommending framework extension examples","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Personalization; Reuse; Extension (predicate logic); Code (set theory); Software engineering; Code reuse; Programming language; Interface (matter); Point (geometry); Software; World Wide Web; Operating system; Set (abstract data type)","score_opus":0.07796341389236565,"score_gpt":0.3504158872168476,"score_spread":0.2724524733244819,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2768062948","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022747938,0.0007167033,0.55303615,0.00057282293,0.00016405775,0.0006679205,0.022052985,0.39222372,0.00781769],"genre_scores_gemma":[0.054521725,0.00041659066,0.88583755,0.0001833262,0.00004347847,0.0006611243,0.04026237,0.012410606,0.0056632934],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986486,0.00018551816,0.00014675023,0.00034906666,0.00057502184,0.00009509124],"domain_scores_gemma":[0.99200845,0.004202639,0.00089844514,0.0010702297,0.0015316107,0.0002886241],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022983286,0.0023642585,0.00090183766,0.012892462,0.00095253246,0.0018618851,0.0024339918,0.0016039547,0.015098452],"category_scores_gemma":[0.020355655,0.0013354446,0.0013891251,0.0038463697,0.00053127215,0.003598691,0.0024332695,0.001495515,0.00816815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038014393,0.00034015588,0.016383173,0.0015320649,0.00013181989,0.0008627182,0.0012258238,0.007361556,0.011227768,0.009415376,0.28219247,0.6689469],"study_design_scores_gemma":[0.00024538717,0.00036667023,0.01604723,0.0008096051,0.00017619121,0.0026447657,0.0015780544,0.48142064,0.039882515,0.03339966,0.42312846,0.00030089848],"about_ca_topic_score_codex":0.0056831795,"about_ca_topic_score_gemma":0.013333576,"teacher_disagreement_score":0.015098452,"about_ca_system_score_codex":0.0007232481,"about_ca_system_score_gemma":0.0019664976,"threshold_uncertainty_score":0.050509334},"labels":[],"label_agreement":null},{"id":"W2768251838","doi":"10.17760/d20467254","title":"Interactive synthesis of code-level security rules","year":2017,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Science North","funders":"","keywords":"Computer science; Software security assurance; Programming language; Software; Code (set theory); Security bug; Process (computing); Software bug; Software engineering; Software development; Software development process; Theoretical computer science; Computer security; Information security; Set (abstract data type)","score_opus":0.02898663466085159,"score_gpt":0.3213096837392981,"score_spread":0.2923230490784465,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2768251838","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06906377,0.000082707425,0.9175332,0.00014361319,0.00004011709,0.00029155996,0.0003551709,0.005238986,0.0072508263],"genre_scores_gemma":[0.4146082,0.00015892384,0.5802634,0.000099055294,0.000015713997,0.00034005672,0.00077531073,0.0005392988,0.0032000595],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992169,0.00017635667,0.000065123415,0.0002140389,0.00025616272,0.00007135239],"domain_scores_gemma":[0.9952076,0.0035154964,0.00021093007,0.00040854508,0.0005935776,0.00006385177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010315336,0.0009304406,0.00040812092,0.00087312557,0.00038746724,0.0009528329,0.0008929707,0.0006336793,0.005393186],"category_scores_gemma":[0.0075874887,0.0004457669,0.00085413194,0.00029470303,0.00068816886,0.0007128127,0.0008316058,0.0007929651,0.0009226333],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003085367,0.00032248726,0.005217579,0.00074885273,0.000098212506,0.0007264185,0.0010649059,0.48312747,0.101439804,0.030101417,0.002828272,0.37401608],"study_design_scores_gemma":[0.00005438205,0.00013119774,0.000435502,0.000056330235,0.00005290681,0.00012995173,0.000103898754,0.90104884,0.08049472,0.011127857,0.0063429377,0.000021483087],"about_ca_topic_score_codex":0.0023631451,"about_ca_topic_score_gemma":0.003584042,"teacher_disagreement_score":0.005393186,"about_ca_system_score_codex":0.0008383428,"about_ca_system_score_gemma":0.0012651839,"threshold_uncertainty_score":0.018041968},"labels":[],"label_agreement":null},{"id":"W2770367167","doi":"10.1002/smr.1916","title":"Model refactoring by example: A multi‐objective search based software engineering approach","year":2017,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"National Science Foundation of Sri Lanka; Qatar Foundation; Qatar National Research Fund; National Science Foundation","keywords":"Code refactoring; Correctness; Computer science; Set (abstract data type); Programming language; Metamodeling; Sorting; Similarity (geometry); Software; Software engineering; Artificial intelligence","score_opus":0.04576571935716419,"score_gpt":0.2955329575172809,"score_spread":0.24976723816011673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2770367167","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09024694,0.00070396037,0.90338486,0.0005263003,0.0000362621,0.00024553965,0.0001360522,0.00065068575,0.0040694536],"genre_scores_gemma":[0.43141627,0.00029662854,0.56560236,0.00018612412,0.000027006323,0.00036526332,0.00028627142,0.000104134495,0.0017158997],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990903,0.000472746,0.000053984575,0.00011813478,0.00019823929,0.00006670257],"domain_scores_gemma":[0.99746275,0.0018740731,0.000195607,0.00010332124,0.00029170106,0.00007260564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019058073,0.0013369005,0.0011426655,0.003201231,0.0004643321,0.000799608,0.0017688384,0.0016616716,0.0026244123],"category_scores_gemma":[0.004253436,0.00067617756,0.0010897571,0.0016159844,0.0005880744,0.00073630956,0.0010595643,0.0008291131,0.00025218553],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004387069,0.00012621614,0.0007780368,0.00010160435,0.000064765234,0.000069366164,0.000058753754,0.95219547,0.0009892454,0.0025887007,0.00051852944,0.042465366],"study_design_scores_gemma":[0.000014794832,0.000034336612,0.00009687317,0.000012203741,0.00001618603,0.00001456125,0.00001556261,0.9981311,0.00024770387,0.0011847442,0.00022807103,0.0000038673174],"about_ca_topic_score_codex":0.006307756,"about_ca_topic_score_gemma":0.007320949,"teacher_disagreement_score":0.006307756,"about_ca_system_score_codex":0.00089588447,"about_ca_system_score_gemma":0.0012215304,"threshold_uncertainty_score":0.012542069},"labels":[],"label_agreement":null},{"id":"W2770939924","doi":"10.1109/sate.2017.16","title":"A Systematic Mapping Study of Quality Assessment Models for Software Products","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Software quality; Quality (philosophy); Metric (unit); Software; Software quality control; Quality assessment; Software metric; Software engineering; Data mining; Reliability engineering; Software development; Engineering; Operations management","score_opus":0.14898901860295502,"score_gpt":0.3874003896231429,"score_spread":0.23841137102018786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2770939924","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11199152,0.69936484,0.13141073,0.006475269,0.0010779037,0.018452123,0.0060070287,0.00037830306,0.02484229],"genre_scores_gemma":[0.5527497,0.26416606,0.15785956,0.0015887278,0.00022436159,0.018093904,0.003456687,0.00011051117,0.0017505396],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.94857854,0.027723663,0.010336858,0.0031357722,0.009734303,0.00049086235],"domain_scores_gemma":[0.7884117,0.15515071,0.01710459,0.00944255,0.02911433,0.0007761639],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.053149622,0.0017700098,0.0027370865,0.042658612,0.0015181856,0.0044717775,0.0017986267,0.0016890906,0.0031207188],"category_scores_gemma":[0.21172869,0.00092930614,0.005297503,0.027239691,0.0014171002,0.0077116564,0.003233056,0.0016654007,0.0004246937],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036025228,0.00042675398,0.01909366,0.29617047,0.0061080996,0.0006682826,0.011443237,0.0036510644,0.0017969249,0.03765868,0.006222966,0.61639965],"study_design_scores_gemma":[0.00072597794,0.0029118513,0.052459884,0.6263645,0.0391706,0.0020490072,0.024421774,0.011688951,0.006862555,0.053430423,0.179467,0.00044744526],"about_ca_topic_score_codex":0.006050624,"about_ca_topic_score_gemma":0.007873979,"teacher_disagreement_score":0.94685036,"about_ca_system_score_codex":0.008453489,"about_ca_system_score_gemma":0.02713598,"threshold_uncertainty_score":0.28108543},"labels":[],"label_agreement":null},{"id":"W2771158352","doi":"10.1002/smr.1925","title":"Evaluating Pred(<i>p</i>) and standardized accuracy criteria in software development effort estimation","year":2017,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Measure (data warehouse); Consistency (knowledge bases); Software; Computer science; Estimation; Statistics; Data mining; Mathematics; Artificial intelligence","score_opus":0.039903469910285164,"score_gpt":0.37960511546840736,"score_spread":0.3397016455581222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2771158352","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7859913,0.0015530476,0.20731747,0.00032569823,0.00006079452,0.00023049011,0.0006526311,0.00060228264,0.003266205],"genre_scores_gemma":[0.9632674,0.00012405086,0.035852466,0.000029040433,0.00002036845,0.00012382475,0.0004231877,0.000032189066,0.00012753089],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9480471,0.028098559,0.0069154683,0.0035498368,0.012384854,0.001004196],"domain_scores_gemma":[0.5195221,0.38988358,0.045485683,0.01504744,0.027840264,0.0022209182],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.053780083,0.0009705195,0.0010052737,0.008352847,0.0005229658,0.0037353458,0.001104202,0.00188633,0.00055774825],"category_scores_gemma":[0.274953,0.0004324647,0.0009866059,0.0073404633,0.0015725319,0.003523544,0.0023012059,0.0012460717,0.00021624654],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006416901,0.00030491344,0.7570458,0.0004431107,0.0010008181,0.00012878775,0.0005981181,0.08456271,0.001742524,0.0047947774,0.00096983847,0.14776672],"study_design_scores_gemma":[0.00007707394,0.0017740881,0.3541366,0.00031872815,0.0003635363,0.00038727224,0.0008383669,0.62097085,0.008272384,0.011166234,0.0015137932,0.0001810358],"about_ca_topic_score_codex":0.0037904363,"about_ca_topic_score_gemma":0.0025404042,"teacher_disagreement_score":0.053780083,"about_ca_system_score_codex":0.001165883,"about_ca_system_score_gemma":0.001952768,"threshold_uncertainty_score":0.28441966},"labels":[],"label_agreement":null},{"id":"W2771280031","doi":"10.1109/sdpc.2017.142","title":"A New Method for Finding Modules of Fault Trees","year":2017,"lang":"en","type":"article","venue":"2017 International Conference on Sensing, Diagnostics, Prognostics, and Control (SDPC)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Natural Science Foundation of China","keywords":"Fault tree analysis; Computer science; Graph; Degree (music); Fault coverage; Tree (set theory); Time complexity; Graph theory; Algorithm; Theoretical computer science; Mathematics; Reliability engineering; Engineering","score_opus":0.060563124906919404,"score_gpt":0.35874561262547583,"score_spread":0.29818248771855643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2771280031","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0023275625,0.00021528096,0.99602425,0.000033681164,0.0000466447,0.00005076569,0.00008533333,0.00053184846,0.0006845139],"genre_scores_gemma":[0.047741767,0.00031667852,0.9480726,0.000065206004,0.000061417515,0.00015835442,0.000423985,0.00017952517,0.0029805081],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99923897,0.00008871307,0.000060211114,0.00023952805,0.0003088233,0.00006381773],"domain_scores_gemma":[0.99922526,0.00020677804,0.00007603233,0.00010739552,0.00033742507,0.00004708334],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00043743948,0.0011169752,0.0011280145,0.0037738874,0.0008817183,0.001009258,0.0015420539,0.0009968414,0.004114989],"category_scores_gemma":[0.001953612,0.0005232253,0.0013240025,0.0023519325,0.0006056178,0.0019602645,0.0012165287,0.00091285416,0.0012680443],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013000297,0.00007004585,0.0020962586,0.0004917452,0.000108583525,0.0002633159,0.00031056584,0.04502951,0.036993057,0.024860756,0.008037543,0.8816086],"study_design_scores_gemma":[0.000071954,0.00019048837,0.0019367727,0.000080816775,0.00012967044,0.0013467189,0.00014731991,0.9015317,0.02254718,0.0329898,0.038921274,0.00010639406],"about_ca_topic_score_codex":0.002962336,"about_ca_topic_score_gemma":0.0030126497,"teacher_disagreement_score":0.004114989,"about_ca_system_score_codex":0.00045592478,"about_ca_system_score_gemma":0.0011693882,"threshold_uncertainty_score":0.013765991},"labels":[],"label_agreement":null},{"id":"W2771651905","doi":"10.1109/esem.2017.46","title":"Which Version Should Be Released to App Store?","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates - Technology Futures","keywords":"Mobile apps; Computer science; Analytics; Predictive analytics; Random forest; Open source; Analogical reasoning; App store; Smartphone app; Mobile device; Data science; Artificial intelligence; World Wide Web","score_opus":0.044740506632467245,"score_gpt":0.3137634228131647,"score_spread":0.26902291618069746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2771651905","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9625276,0.0014510788,0.020806683,0.0016063655,0.00009563868,0.00026183928,0.0026850044,0.0013851409,0.009180621],"genre_scores_gemma":[0.9816001,0.0005600324,0.013721652,0.00020470751,0.00007105741,0.000072560266,0.0016796608,0.00021289277,0.0018774496],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99592245,0.0006118721,0.0005516497,0.0011054733,0.0015566072,0.00025188076],"domain_scores_gemma":[0.94187456,0.035486244,0.012251687,0.0040272577,0.005232548,0.0011276142],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004451111,0.0005243362,0.00036828071,0.0021857147,0.0009541388,0.0021909112,0.000849308,0.0008765307,0.0032799777],"category_scores_gemma":[0.05397865,0.0003794902,0.00064412225,0.0017311077,0.00092720607,0.0044372627,0.00069795264,0.0010068791,0.0011121358],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015512825,0.00041435778,0.5587476,0.0014759616,0.000202641,0.0030961258,0.0046868916,0.008817203,0.016143002,0.0067388103,0.010015933,0.3881102],"study_design_scores_gemma":[0.00012982047,0.0010481454,0.7545453,0.0011219254,0.0004288092,0.0071008466,0.004283741,0.12261503,0.02795261,0.018685361,0.061817333,0.0002711172],"about_ca_topic_score_codex":0.0047555086,"about_ca_topic_score_gemma":0.00418431,"teacher_disagreement_score":0.0047555086,"about_ca_system_score_codex":0.00097456784,"about_ca_system_score_gemma":0.00095946575,"threshold_uncertainty_score":0.02354002},"labels":[],"label_agreement":null},{"id":"W2771968043","doi":"10.1109/esem.2017.66","title":"Beyond Boxes and Lines: Creating and Empirically Evaluating Alternative Visualizations for Requirements Conceptual Models","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Diagrammatic reasoning; Visualization; Modeling language; Notation; Semantics (computer science); Domain (mathematical analysis); Domain model; Conceptual model; Domain-specific language; Software engineering; Human–computer interaction; Data science; Programming language; Natural language processing; Artificial intelligence; Software; Domain knowledge; Database; Linguistics","score_opus":0.17731079720641843,"score_gpt":0.44435911905114006,"score_spread":0.2670483218447216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2771968043","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7751116,0.0011123653,0.2036071,0.0019476556,0.0001712838,0.005665599,0.0010856837,0.0017513536,0.009547314],"genre_scores_gemma":[0.649613,0.0003163433,0.34390095,0.00024199222,0.00003036723,0.004267196,0.0009724777,0.00021616562,0.00044149076],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.91673964,0.06985166,0.005549142,0.0026009781,0.0046162107,0.0006423178],"domain_scores_gemma":[0.2988824,0.6461046,0.013604283,0.026601978,0.012980413,0.0018264484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.088704705,0.0018569432,0.00095290504,0.0045344434,0.0017007726,0.0071783904,0.0034182945,0.0033553503,0.0047207656],"category_scores_gemma":[0.43191355,0.0011413068,0.0015925907,0.0035556061,0.0040790215,0.013908029,0.0074842568,0.004008709,0.0005718629],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009871775,0.011211712,0.088161305,0.010952255,0.00102518,0.0010669654,0.12917241,0.11757786,0.014096723,0.13178799,0.012273602,0.4728023],"study_design_scores_gemma":[0.0055663404,0.014846912,0.05175714,0.0068272525,0.0014911187,0.0007764019,0.056820925,0.6660914,0.023859674,0.13226822,0.038674586,0.0010200482],"about_ca_topic_score_codex":0.0026252512,"about_ca_topic_score_gemma":0.0048991414,"teacher_disagreement_score":0.088704705,"about_ca_system_score_codex":0.003843935,"about_ca_system_score_gemma":0.0020528329,"threshold_uncertainty_score":0.46912092},"labels":[],"label_agreement":null},{"id":"W2771971885","doi":"10.1145/3196398.3196400","title":"Enriched event streams","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Event (particle physics); Context (archaeology); STREAMS; Empirical research; Code (set theory); Source code; Process (computing); Development environment; Human–computer interaction; World Wide Web; Software engineering; Programming language; Operating system","score_opus":0.014316942736155988,"score_gpt":0.2818035671668636,"score_spread":0.26748662443070764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2771971885","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29285702,0.0032263352,0.50667745,0.0010700292,0.0007440246,0.002528886,0.15302207,0.023809921,0.016064234],"genre_scores_gemma":[0.61065316,0.0016419329,0.23672584,0.00058817567,0.00052575755,0.0024575396,0.14042903,0.0016835213,0.0052950964],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963381,0.0007356445,0.00036815146,0.0009534148,0.0013569172,0.0002477983],"domain_scores_gemma":[0.98526764,0.007949295,0.0014576071,0.0016145183,0.0031767974,0.00053424394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019314894,0.0013745839,0.0008889993,0.0031941691,0.0005807057,0.001920949,0.00092903257,0.0008152882,0.0031957426],"category_scores_gemma":[0.019590152,0.0004301657,0.0004859653,0.003604021,0.00037132844,0.0019611975,0.0017684806,0.0011398526,0.0017710937],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003935828,0.0009373501,0.20917545,0.0044241003,0.00084211957,0.0037074117,0.00609724,0.024783844,0.035835933,0.017255608,0.08852301,0.6044821],"study_design_scores_gemma":[0.00059552677,0.0009424911,0.2962952,0.0010741042,0.0007480433,0.004036141,0.003781397,0.20974225,0.06334638,0.08614318,0.33282694,0.00046834032],"about_ca_topic_score_codex":0.0025379825,"about_ca_topic_score_gemma":0.004257895,"teacher_disagreement_score":0.0031957426,"about_ca_system_score_codex":0.00047394822,"about_ca_system_score_gemma":0.00079628825,"threshold_uncertainty_score":0.010690808},"labels":[],"label_agreement":null},{"id":"W2772402907","doi":"10.1007/s00500-017-2945-4","title":"On the value of parameter tuning in heterogeneous ensembles effort estimation","year":2017,"lang":"en","type":"article","venue":"Soft Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Particle swarm optimization; Computer science; Software; Hyperparameter optimization; Grid; Machine learning; Data mining; Perceptron; Support vector machine; Estimation; Artificial intelligence; Artificial neural network; Mathematics; Engineering","score_opus":0.026627897201802822,"score_gpt":0.29207575307502864,"score_spread":0.26544785587322584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2772402907","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12238599,0.002847027,0.8650923,0.0029084282,0.00019466688,0.00013073742,0.0002064421,0.0006034627,0.0056309155],"genre_scores_gemma":[0.9024648,0.00058824517,0.09462709,0.00043055083,0.00027489473,0.00011468261,0.00024349937,0.00017214191,0.0010840746],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9896158,0.0076133744,0.00038937302,0.0010741756,0.00087006774,0.00043721264],"domain_scores_gemma":[0.85771525,0.13219047,0.0015597184,0.005370711,0.00232691,0.0008368715],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0213497,0.0016894041,0.0024950954,0.0012761919,0.0008940166,0.0027165185,0.002567626,0.0038605,0.0017516738],"category_scores_gemma":[0.1618317,0.0009394403,0.00095069467,0.0010449796,0.0018303432,0.005293003,0.0035220191,0.0038596326,0.0003577382],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011147219,0.00029777767,0.008893082,0.00027197754,0.00042749898,0.00016264498,0.00027230143,0.8180856,0.0026688688,0.033168428,0.0025204911,0.13211654],"study_design_scores_gemma":[0.000025052648,0.00004257354,0.0006340222,0.00005300617,0.0000394291,0.000025526724,0.000034399036,0.98608553,0.0005845821,0.01227929,0.00018329776,0.00001326497],"about_ca_topic_score_codex":0.0030155412,"about_ca_topic_score_gemma":0.0018850913,"teacher_disagreement_score":0.0213497,"about_ca_system_score_codex":0.00093526015,"about_ca_system_score_gemma":0.001448454,"threshold_uncertainty_score":0.11290932},"labels":[],"label_agreement":null},{"id":"W2772825062","doi":"10.1109/esem.2017.25","title":"An Ontology-Based Approach to Automate Tagging of Software Artifacts","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Commit; Ontology; Context (archaeology); Software; Software security assurance; Prioritization; Information retrieval; Software engineering; Database; Information security; Computer security; Programming language","score_opus":0.03448545336031505,"score_gpt":0.3097950461178106,"score_spread":0.27530959275749556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2772825062","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059371255,0.00012127187,0.9837114,0.0005693037,0.00008426293,0.00054540974,0.0016812343,0.0048060436,0.00254399],"genre_scores_gemma":[0.03216361,0.00016439406,0.9612188,0.00016712693,0.00003186444,0.00031611425,0.0034254452,0.0003146813,0.0021980228],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964994,0.00075604254,0.00054573506,0.00080611784,0.0012044155,0.0001882979],"domain_scores_gemma":[0.99184424,0.0026225871,0.0010920742,0.001965848,0.0021835973,0.00029162652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034993158,0.0011122222,0.00069403223,0.0076443795,0.0019234915,0.0033246675,0.0016405915,0.0015542752,0.002455198],"category_scores_gemma":[0.011021193,0.0007535927,0.0019207917,0.0049984744,0.0012397742,0.005071618,0.0030418076,0.0030482307,0.0029190846],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013464082,0.00089680677,0.009732764,0.0013940575,0.0002648909,0.000936878,0.0060658273,0.009381857,0.07762008,0.05973796,0.031093588,0.8027407],"study_design_scores_gemma":[0.00008956696,0.00024252018,0.01804896,0.0009846655,0.0004468749,0.0019320714,0.0038534633,0.42398557,0.07866546,0.15341075,0.31796768,0.00037243782],"about_ca_topic_score_codex":0.010700156,"about_ca_topic_score_gemma":0.01992915,"teacher_disagreement_score":0.010700156,"about_ca_system_score_codex":0.0019857604,"about_ca_system_score_gemma":0.0052785967,"threshold_uncertainty_score":0.021275759},"labels":[],"label_agreement":null},{"id":"W2780123483","doi":"10.22495/rgc7i4c2art2","title":"The use of generalised audit software by internal audit functions in a developing country: A maturity level assessment","year":2017,"lang":"en","type":"article","venue":"Risk Governance and Control Financial Markets & Institutions","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Maturity (psychological); Internal audit; Audit; Accounting; Business; Capability Maturity Model; Benchmark (surveying); Information technology audit; Software; Order (exchange); Perspective (graphical); Computer science; Finance; Geography; Joint audit; Political science","score_opus":0.02875939799489748,"score_gpt":0.27080540098516903,"score_spread":0.24204600299027154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2780123483","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9787329,0.0015533036,0.008442278,0.0011854797,0.0000145991535,0.00016849182,0.00018921673,0.00009622934,0.009617466],"genre_scores_gemma":[0.9923964,0.00084668765,0.005632418,0.000057200326,0.0000049250525,0.000038830825,0.000114856346,0.000017108187,0.0008915848],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99110246,0.0035274697,0.0010033462,0.0006189155,0.0029547156,0.0007931559],"domain_scores_gemma":[0.9260313,0.025889747,0.020726383,0.0066107856,0.018054854,0.0026870386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016826428,0.00024453804,0.00023854809,0.0047765956,0.00073716266,0.0036916586,0.0005806407,0.0004465594,0.0012752394],"category_scores_gemma":[0.05560771,0.00034817497,0.0003085992,0.0042334083,0.0015037976,0.002999919,0.0025720738,0.00083592144,0.0003374519],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020034474,0.00012345964,0.6033982,0.000597512,0.000053568237,0.0005031632,0.038094997,0.002058232,0.0065059196,0.008877071,0.00097031146,0.33861727],"study_design_scores_gemma":[0.000022385682,0.00093970314,0.8949545,0.0020560934,0.000080504535,0.0016698117,0.042924605,0.0057950253,0.00823876,0.004941062,0.038220495,0.00015706234],"about_ca_topic_score_codex":0.006200343,"about_ca_topic_score_gemma":0.0076562646,"teacher_disagreement_score":0.016826428,"about_ca_system_score_codex":0.0026320808,"about_ca_system_score_gemma":0.004784629,"threshold_uncertainty_score":0.08898777},"labels":[],"label_agreement":null},{"id":"W2781233888","doi":"10.7763/ijcte.2017.v9.1160","title":"Review on JPA Based ORM Data Persistence Framework","year":2017,"lang":"en","type":"article","venue":"International Journal of Computer Theory and Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Persistence (discontinuity); Persistent data structure; Programming language","score_opus":0.046680216497100704,"score_gpt":0.3171926362971748,"score_spread":0.2705124198000741,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2781233888","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00065363944,0.97995126,0.0075887707,0.0012915897,0.00074614835,0.00005189785,0.00019194362,0.00037136566,0.009153485],"genre_scores_gemma":[0.0054361382,0.97132695,0.0137760285,0.0017411658,0.0007923305,0.00013126506,0.0005975229,0.00017747378,0.0060210135],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990374,0.0001342664,0.00013211108,0.00015807012,0.00046105246,0.0000772438],"domain_scores_gemma":[0.9982502,0.00085057155,0.0002141218,0.00009533929,0.00050566683,0.00008420341],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013964808,0.0006576674,0.00083517464,0.0031707317,0.0005050958,0.0016261372,0.0020212114,0.0013539129,0.0064755655],"category_scores_gemma":[0.0034846545,0.0005073729,0.0010615587,0.0036373623,0.00052180537,0.003680407,0.0009538866,0.0014845047,0.0036916162],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000091525646,0.000069675196,0.00035021748,0.023167154,0.00007050189,0.00018645232,0.0002048925,0.00076961156,0.0035812897,0.020735169,0.07663713,0.8741364],"study_design_scores_gemma":[0.0000059334016,0.00003322576,0.00028683167,0.0022339218,0.000057373603,0.0003342879,0.000037190224,0.00017699572,0.000807157,0.0014994037,0.99451077,0.000016890652],"about_ca_topic_score_codex":0.002836271,"about_ca_topic_score_gemma":0.0018863202,"teacher_disagreement_score":0.0064755655,"about_ca_system_score_codex":0.0009140248,"about_ca_system_score_gemma":0.002908366,"threshold_uncertainty_score":0.02166295},"labels":[],"label_agreement":null},{"id":"W2781488941","doi":"10.1007/978-981-10-7796-8_2","title":"An Empirical Study of the Software Development Process, Including Its Requirements Engineering, at Very Large Organization: How to Use Data Mining in Such a Study","year":2018,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Process (computing); Computer science; Software; Empirical research; Focus (optics); Data science; Software engineering","score_opus":0.15903600506905624,"score_gpt":0.379249727736557,"score_spread":0.22021372266750078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2781488941","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69154483,0.0153593095,0.15392418,0.026224883,0.0004090801,0.0005485756,0.0019518624,0.00042574614,0.10961164],"genre_scores_gemma":[0.88412315,0.008483376,0.090808064,0.001815397,0.00018382036,0.00035977547,0.0011651913,0.00013513664,0.012926202],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9978393,0.0012932925,0.00013497878,0.00015593675,0.0004816676,0.00009488089],"domain_scores_gemma":[0.9531679,0.041301787,0.0015424826,0.0014076759,0.0020413054,0.00053878716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00421525,0.00028694022,0.0002952579,0.0020480014,0.0011392653,0.0020196433,0.0008530022,0.0007527203,0.0026384061],"category_scores_gemma":[0.02440653,0.00031904475,0.00038884915,0.0048504802,0.0013250326,0.0057527036,0.0012703139,0.0016924315,0.00054623303],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001699318,0.0019125827,0.2591496,0.0016166006,0.00014811005,0.0012667582,0.024419231,0.0037132434,0.0033326622,0.12398994,0.039639026,0.5406423],"study_design_scores_gemma":[0.00006653882,0.0012217483,0.55517286,0.0031520978,0.00023414561,0.0022655495,0.06791787,0.043258246,0.00735267,0.18400271,0.13519242,0.00016326556],"about_ca_topic_score_codex":0.004811805,"about_ca_topic_score_gemma":0.0082973605,"teacher_disagreement_score":0.004811805,"about_ca_system_score_codex":0.0011058116,"about_ca_system_score_gemma":0.0017200385,"threshold_uncertainty_score":0.022292674},"labels":[],"label_agreement":null},{"id":"W2781935674","doi":"","title":"Towards a Theory of Technical Debt Ownership: An Exploratory Field Study","year":2017,"lang":"en","type":"article","venue":"Journal of the Association for Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Accountability; Technical debt; Exploratory research; Context (archaeology); Total cost of ownership; Quality (philosophy); Accounting; Debt; Field (mathematics); Qualitative research; Process management; Business; Software development; Marketing; Knowledge management; Software; Computer science; Finance; Political science; Sociology","score_opus":0.0360676942486739,"score_gpt":0.3054190181517246,"score_spread":0.2693513239030507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2781935674","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9489347,0.0006193013,0.019314954,0.004184225,0.00003419727,0.000946557,0.00015158638,0.000012361934,0.025802212],"genre_scores_gemma":[0.99229586,0.0003531035,0.0053294366,0.00043764815,0.00001698531,0.0006553516,0.000046157034,0.0000054996385,0.00085995893],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.98140895,0.014482733,0.00052961236,0.0011042581,0.0014381509,0.0010362209],"domain_scores_gemma":[0.93173885,0.05923065,0.0040155794,0.001435407,0.0024082817,0.0011712256],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023568356,0.00043468378,0.00061321026,0.0043104943,0.005119346,0.006447751,0.002041719,0.0018839865,0.005930063],"category_scores_gemma":[0.028519494,0.0006636923,0.00048192203,0.0040836167,0.013163319,0.0105654225,0.0048463475,0.0030727,0.00035722615],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012809678,0.002464625,0.09415079,0.00078967866,0.00002441608,0.0018669807,0.7047438,0.0007615305,0.0011580965,0.1511397,0.0017450916,0.041027237],"study_design_scores_gemma":[0.00005579722,0.00036648553,0.028610997,0.00091697305,0.000017546678,0.00042636762,0.9156179,0.0040531545,0.0003513116,0.038979117,0.010574036,0.000030311723],"about_ca_topic_score_codex":0.0038254918,"about_ca_topic_score_gemma":0.003773909,"teacher_disagreement_score":0.023568356,"about_ca_system_score_codex":0.0057425336,"about_ca_system_score_gemma":0.005776017,"threshold_uncertainty_score":0.12464291},"labels":[],"label_agreement":null},{"id":"W2783244927","doi":"10.4018/978-1-60566-026-4.ch483","title":"Pattern-Oriented Use Case Modeling","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software deployment; Process (computing); Software engineering; Presentation (obstetrics); Software; Software development; Quality (philosophy); Use Case Points; Class (philosophy); Software development process; Data science; Artificial intelligence; Programming language","score_opus":0.03368052633234563,"score_gpt":0.26508812491221156,"score_spread":0.23140759857986593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2783244927","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077909557,0.00024065042,0.9714575,0.00074889924,0.00003516305,0.0025357811,0.0009569938,0.00062601676,0.015608066],"genre_scores_gemma":[0.06457744,0.00054942,0.92376214,0.00011364872,0.000023906889,0.0033040687,0.0016295761,0.00012511926,0.005914806],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9854872,0.0069712056,0.0016760749,0.0015650018,0.0037420855,0.00055846653],"domain_scores_gemma":[0.97970515,0.014198082,0.0013482672,0.002262751,0.0021888607,0.00029684772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011005907,0.0015130665,0.0008633014,0.0068673124,0.0015207409,0.0077043176,0.005746423,0.002790859,0.009007126],"category_scores_gemma":[0.02488352,0.0013320886,0.0029654796,0.006188897,0.0017911709,0.0063989903,0.004135211,0.0020366586,0.0018708014],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014899771,0.00051726954,0.006460883,0.001215193,0.00019734849,0.0029024882,0.009539273,0.08991838,0.0037494635,0.64864373,0.009474137,0.2272328],"study_design_scores_gemma":[0.00012917744,0.0001628285,0.0016915088,0.0008754447,0.0001591228,0.0020951626,0.0028323797,0.4803781,0.0055991258,0.30683693,0.19910741,0.0001328194],"about_ca_topic_score_codex":0.0055984,"about_ca_topic_score_gemma":0.0070318608,"teacher_disagreement_score":0.011005907,"about_ca_system_score_codex":0.0032871922,"about_ca_system_score_gemma":0.0036704317,"threshold_uncertainty_score":0.058205485},"labels":[],"label_agreement":null},{"id":"W2783444782","doi":"10.1109/icdim.2017.8244675","title":"Code style analytics for the automatic setting of formatting rules in IDEs: A solution to the Tabs vs. Spaces Debate","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Disk formatting; Computer science; Source code; Code review; Code (set theory); Source lines of code; Style (visual arts); Analytics; Programming language; Software; Class (philosophy); Static program analysis; KPI-driven code analysis; Software engineering; World Wide Web; Database; Software development; Artificial intelligence; Operating system; Set (abstract data type)","score_opus":0.0330001482133736,"score_gpt":0.3145460822316457,"score_spread":0.2815459340182721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2783444782","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042657096,0.0006023068,0.8976945,0.0011315284,0.00018221677,0.0005553078,0.0049984427,0.049256373,0.002922188],"genre_scores_gemma":[0.15521981,0.0003038323,0.8307685,0.00022361458,0.00012751532,0.000445957,0.0076243845,0.0037478213,0.0015384894],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9880199,0.0037864537,0.0016950743,0.0024222166,0.0036372945,0.00043906982],"domain_scores_gemma":[0.9398841,0.025478523,0.006734549,0.014234278,0.01217642,0.001492066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01051646,0.0019628522,0.0016854729,0.012964668,0.0011365755,0.006414267,0.002447565,0.0015890563,0.0032017238],"category_scores_gemma":[0.068651825,0.001122517,0.0013941155,0.0074996552,0.001176106,0.0070098173,0.003660721,0.0032797735,0.0038083943],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004785595,0.00054412265,0.04236681,0.0006058629,0.00016566197,0.00022672952,0.0023875728,0.0067048203,0.01115146,0.010795881,0.022311281,0.90226114],"study_design_scores_gemma":[0.00025135162,0.0004337412,0.027105557,0.0006466455,0.00011384696,0.0006012346,0.0021259566,0.77368146,0.048474617,0.084466584,0.061772455,0.00032652012],"about_ca_topic_score_codex":0.0020571982,"about_ca_topic_score_gemma":0.0032004898,"teacher_disagreement_score":0.012964668,"about_ca_system_score_codex":0.0009205661,"about_ca_system_score_gemma":0.002014231,"threshold_uncertainty_score":0.055617034},"labels":[],"label_agreement":null},{"id":"W2783794442","doi":"10.1016/j.infsof.2018.01.003","title":"Support vector regression for predicting software enhancement effort","year":2018,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Consejo Nacional de Ciencia y Tecnología; National Research Council Canada; Universidad de Guadalajara","keywords":"Support vector machine; Computer science; Benchmarking; Machine learning; Artificial neural network; Data mining; Software; Sigmoid function; Artificial intelligence; Decision tree; Set (abstract data type); Regression analysis; Kernel (algebra); Radial basis function kernel; Kernel method; Mathematics; Operating system","score_opus":0.010772910341071697,"score_gpt":0.2680281242044023,"score_spread":0.2572552138633306,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2783794442","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84714895,0.00071441755,0.14872685,0.00022422228,0.00007475007,0.000056409892,0.0006561537,0.0013635126,0.0010347111],"genre_scores_gemma":[0.96217984,0.0001770901,0.035083003,0.000016685799,0.000029772276,0.00003701137,0.0007338025,0.000043611795,0.0016991687],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988426,0.00056606624,0.0000657928,0.00015419158,0.0002848288,0.00008640238],"domain_scores_gemma":[0.9859154,0.011068365,0.0009189697,0.00065409625,0.0012147771,0.00022832258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002530915,0.00086784497,0.0005176873,0.0015975116,0.00015672047,0.0004892022,0.00065449515,0.0006786282,0.0014997085],"category_scores_gemma":[0.0172024,0.00020138707,0.00042005393,0.001418291,0.00015476247,0.00096086954,0.0003439657,0.0011623633,0.00085351],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013401789,0.0017866205,0.13045773,0.00021960176,0.00029473729,0.00012734969,0.00014399776,0.2416735,0.0068932134,0.0013586791,0.0042581107,0.61144626],"study_design_scores_gemma":[0.000018149374,0.00023656007,0.013170833,0.000012001406,0.000034378914,0.000019949888,0.000031885545,0.98380154,0.0017566303,0.0006849268,0.00022294743,0.000010286436],"about_ca_topic_score_codex":0.0041677402,"about_ca_topic_score_gemma":0.0033717828,"teacher_disagreement_score":0.0041677402,"about_ca_system_score_codex":0.0002709794,"about_ca_system_score_gemma":0.0005301769,"threshold_uncertainty_score":0.013384938},"labels":[],"label_agreement":null},{"id":"W2783803218","doi":"","title":"Assisting developers towards fault localization by analyzing failure reports","year":2017,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Root cause; Ranking (information retrieval); Source code; Call graph; Context (archaeology); Data mining; Set (abstract data type); Focus (optics); Interdependence; Root cause analysis; Software bug; Software; Information retrieval; Theoretical computer science; Programming language; Reliability engineering; Engineering","score_opus":0.012473060674922955,"score_gpt":0.2480784589620084,"score_spread":0.23560539828708543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2783803218","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.108143166,0.00031466485,0.8731207,0.000546622,0.000024422354,0.00022818292,0.0003682998,0.016150663,0.0011032296],"genre_scores_gemma":[0.42230177,0.00032171464,0.5745534,0.00009586637,0.00004504751,0.000108776905,0.0009331805,0.0005102043,0.0011300319],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.995621,0.0019610885,0.00029827334,0.0007462088,0.0011067287,0.00026665724],"domain_scores_gemma":[0.95311964,0.029668879,0.005689817,0.0050722677,0.005631509,0.0008179307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006959543,0.0029701642,0.0016458413,0.008219499,0.00081538653,0.0032582222,0.0018107056,0.0017516445,0.0020057438],"category_scores_gemma":[0.04993256,0.0011087877,0.00084281137,0.00258737,0.0007488582,0.0029996706,0.0022288593,0.0018281065,0.002875362],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006343776,0.0009729345,0.071517974,0.0006147883,0.00015988386,0.0004811318,0.0015023716,0.076831035,0.035363372,0.0023565514,0.005980654,0.8035849],"study_design_scores_gemma":[0.00007393299,0.00041725504,0.009744263,0.00010422397,0.00014258313,0.000433549,0.00083807076,0.929888,0.044220902,0.00966917,0.0043708193,0.00009716058],"about_ca_topic_score_codex":0.0034458453,"about_ca_topic_score_gemma":0.0064141173,"teacher_disagreement_score":0.008219499,"about_ca_system_score_codex":0.0006272924,"about_ca_system_score_gemma":0.0027735732,"threshold_uncertainty_score":0.036806047},"labels":[],"label_agreement":null},{"id":"W2783864264","doi":"10.5539/mas.v12n2p54","title":"Grey Wolf Algorithm for Requirements Prioritization","year":2018,"lang":"en","type":"article","venue":"Modern Applied Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Prioritization; Viewpoints; Analytic hierarchy process; Computer science; Requirement prioritization; sort; Hierarchy; Software; Process (computing); Order (exchange); Sorting; Operations research; Requirements engineering; Data mining; Algorithm; Mathematics; Management science; Requirements management; Database; Engineering","score_opus":0.030653078199211755,"score_gpt":0.3014379095046339,"score_spread":0.27078483130542214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2783864264","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009571364,0.00038537112,0.9866166,0.00020605646,0.00004143391,0.000120042634,0.00008991627,0.0003477679,0.0026214097],"genre_scores_gemma":[0.32729048,0.00052917254,0.6662127,0.00023233629,0.00005139438,0.00056991656,0.00048924686,0.00012470466,0.004500099],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904615,0.00031038487,0.000076492055,0.00016848366,0.0002827364,0.00011576742],"domain_scores_gemma":[0.9992124,0.00047237877,0.00006958908,0.000035780005,0.0001772098,0.000032681535],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014586943,0.0010362684,0.0013922222,0.0016834802,0.00064729,0.0010281242,0.0010599247,0.0010771413,0.003889249],"category_scores_gemma":[0.0028371369,0.0004838615,0.0010430089,0.0016260002,0.0005143645,0.0009983211,0.0009930832,0.0010686009,0.00055851456],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010524177,0.000094687304,0.0011752148,0.0002635952,0.00011197268,0.00013176694,0.00018121359,0.7426163,0.003515335,0.0150307985,0.0042964392,0.23247743],"study_design_scores_gemma":[0.00002562876,0.00005147136,0.00021776927,0.000018047876,0.00001491883,0.00003207306,0.000035375513,0.98870426,0.00070516847,0.008487019,0.0017004563,0.000007854642],"about_ca_topic_score_codex":0.0067749503,"about_ca_topic_score_gemma":0.004485541,"teacher_disagreement_score":0.0067749503,"about_ca_system_score_codex":0.0008932817,"about_ca_system_score_gemma":0.002281649,"threshold_uncertainty_score":0.013471067},"labels":[],"label_agreement":null},{"id":"W2784023338","doi":"10.1145/3168365.3168376","title":"VarXplorer","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Feature (linguistics); Focus (optics); Representation (politics); Visualization; Feature model; Pairwise comparison; Theoretical computer science; Human–computer interaction; Artificial intelligence; Programming language; Software","score_opus":0.024175106323710064,"score_gpt":0.27493759316867733,"score_spread":0.25076248684496727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2784023338","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005068716,0.00039260823,0.35519835,0.00050731824,0.00018391905,0.0002866389,0.015768642,0.6045847,0.018009074],"genre_scores_gemma":[0.12323866,0.0014458401,0.43336344,0.0023839385,0.00021859346,0.0019519643,0.06245257,0.28722998,0.087715045],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998697,0.00024432968,0.00015683024,0.00029147157,0.00047136712,0.0001391059],"domain_scores_gemma":[0.9942702,0.0038915607,0.000339004,0.0008250248,0.0005193079,0.00015496819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024306574,0.0027665996,0.001055819,0.0019831406,0.000696754,0.0030763159,0.0043833144,0.0027521844,0.094815716],"category_scores_gemma":[0.009630777,0.0020596618,0.002040712,0.0010197787,0.0008897595,0.0058362386,0.0034065247,0.003105794,0.031062786],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016530809,0.00031320687,0.006931685,0.0033619916,0.00020914097,0.0027190312,0.0016065296,0.009850595,0.03210134,0.050840113,0.5981593,0.292254],"study_design_scores_gemma":[0.0005305838,0.00026918156,0.0047106855,0.00074953074,0.00009946041,0.0030195087,0.00028023493,0.06636277,0.052237503,0.04740292,0.8239931,0.0003444888],"about_ca_topic_score_codex":0.0021284078,"about_ca_topic_score_gemma":0.0028373608,"teacher_disagreement_score":0.094815716,"about_ca_system_score_codex":0.00080149336,"about_ca_system_score_gemma":0.001087547,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2785953012","doi":"10.1142/s0218194017400083","title":"Comparing Software Bugs in Clone and Non-clone Code: An Empirical Study","year":2017,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Software maintenance; Cloning (programming); Commit; Code (set theory); Code refactoring; Computer science; Software bug; Programming language; Source code; Software; Software engineering; Software system; Biology; Database; Genetics; Set (abstract data type)","score_opus":0.030350166117151002,"score_gpt":0.32793527347887125,"score_spread":0.2975851073617202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2785953012","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99933785,0.00008363566,0.00030623886,0.00002230449,0.0000018314927,0.000028036833,0.000049814647,0.0000063135567,0.00016394036],"genre_scores_gemma":[0.9990169,0.00006930323,0.0005518623,0.000021808473,0.0000066909074,0.00004106244,0.00015031338,0.000009326235,0.00013280033],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9909072,0.0034684334,0.0010220318,0.0013105338,0.00280885,0.0004829173],"domain_scores_gemma":[0.7626105,0.16470237,0.044240475,0.007459433,0.016998757,0.003988464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00834256,0.00035861935,0.00043009972,0.0037291208,0.0006303582,0.0010021968,0.00092534255,0.0010025516,0.0016089841],"category_scores_gemma":[0.082886085,0.00037191887,0.00036310044,0.0022642952,0.0019404967,0.002397187,0.0014387071,0.0011730124,0.0003554857],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034530886,0.0009842212,0.97322536,0.00021459331,0.00008297559,0.0004992602,0.008013858,0.00021634197,0.0010658565,0.00019593092,0.0003307574,0.014825371],"study_design_scores_gemma":[0.000033721437,0.0011948398,0.98656785,0.00006403539,0.000047418718,0.0010460342,0.007846476,0.0013700722,0.0007129358,0.00018983828,0.0009020259,0.000024739315],"about_ca_topic_score_codex":0.0014976149,"about_ca_topic_score_gemma":0.0022122273,"teacher_disagreement_score":0.00834256,"about_ca_system_score_codex":0.00063997624,"about_ca_system_score_gemma":0.0005524444,"threshold_uncertainty_score":0.044120252},"labels":[],"label_agreement":null},{"id":"W2786424616","doi":"10.1007/s10664-018-9595-8","title":"Studying software logging using topic models","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":95,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Queen's University","funders":"","keywords":"Logging; Computer science; Software; Software engineering; Operating system; Forestry; Geography","score_opus":0.0777747160628274,"score_gpt":0.31632667881580295,"score_spread":0.23855196275297555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2786424616","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.755968,0.0026501513,0.22799452,0.0024793898,0.00007404296,0.00011724861,0.0002920903,0.00048303817,0.009941475],"genre_scores_gemma":[0.9888076,0.0007088084,0.008644212,0.000068406065,0.00009874165,0.00005865646,0.00026089026,0.00008941415,0.0012631406],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9963007,0.0026723633,0.00012356688,0.00037212027,0.0003120316,0.00021930184],"domain_scores_gemma":[0.8847153,0.10646708,0.0029517377,0.0032276143,0.0016903467,0.0009480038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007309626,0.0007300656,0.00083487347,0.0025919345,0.0010110938,0.0040203393,0.001142146,0.0015939303,0.0032530588],"category_scores_gemma":[0.060726386,0.0007223399,0.00089385477,0.00410205,0.0010399777,0.01049256,0.001217949,0.0024139315,0.00065260875],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013052016,0.0025179693,0.2996075,0.00107492,0.00082719815,0.0006077637,0.01995619,0.11586348,0.0070619653,0.22078556,0.01067179,0.31972048],"study_design_scores_gemma":[0.000118958225,0.0005209794,0.0523946,0.0002086508,0.00033315222,0.00054835813,0.00903551,0.6887113,0.0032829496,0.23535807,0.009396489,0.00009102552],"about_ca_topic_score_codex":0.0031644837,"about_ca_topic_score_gemma":0.0032836339,"teacher_disagreement_score":0.007309626,"about_ca_system_score_codex":0.0013039203,"about_ca_system_score_gemma":0.0009043756,"threshold_uncertainty_score":0.038657486},"labels":[],"label_agreement":null},{"id":"W2786656115","doi":"10.22215/etd/2017-11954","title":"An investigation of software vulnerabilities in open source software projects using data from publicly-available online sources.","year":2017,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Workflow; Software project management; Commit; Software peer review; Software; Software engineering; Data science; Scripting language; World Wide Web; Software quality; Secure coding; Database; Software development; Computer security; Software construction; Software security assurance; Information security; Operating system","score_opus":0.13382109037581552,"score_gpt":0.35852252454556854,"score_spread":0.22470143416975302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2786656115","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9745848,0.00078993884,0.004299251,0.0006847205,0.000019100293,0.0003729187,0.015594351,0.000049125963,0.003605869],"genre_scores_gemma":[0.954822,0.0011887439,0.015198387,0.00015580907,0.000039102746,0.0011266285,0.025767336,0.000063244646,0.0016387362],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99149686,0.0018754213,0.0010050251,0.0009859339,0.004129572,0.0005071945],"domain_scores_gemma":[0.90995413,0.03973424,0.027999347,0.0066356747,0.01389299,0.0017835711],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009992364,0.00032647283,0.00030848625,0.015097836,0.0008813689,0.0021397246,0.0005608067,0.00069381035,0.0011210648],"category_scores_gemma":[0.059813466,0.00035487945,0.0005019588,0.020949919,0.0008598607,0.003438089,0.0034225043,0.0011106911,0.000698435],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006242363,0.00021060443,0.93691653,0.00045950167,0.00010980332,0.000207753,0.010227677,0.00028815906,0.00083347084,0.0022479002,0.0034015696,0.045034617],"study_design_scores_gemma":[0.000009682046,0.00009837823,0.97056645,0.00047615115,0.00009294849,0.0003433921,0.010736706,0.0013985507,0.001489177,0.0011476068,0.013602895,0.00003806066],"about_ca_topic_score_codex":0.008848547,"about_ca_topic_score_gemma":0.013553744,"teacher_disagreement_score":0.015097836,"about_ca_system_score_codex":0.0011234682,"about_ca_system_score_gemma":0.003224772,"threshold_uncertainty_score":0.0528453},"labels":[],"label_agreement":null},{"id":"W2790826992","doi":"10.1177/0741088317748590","title":"Coding for Language Complexity: The Interplay Among Methodological Commitments, Tools, and Workflow in Writing Research","year":2018,"lang":"en","type":"article","venue":"Written Communication","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Coding (social sciences); Computer science; Workflow; Knowledge management; Data science; Sociology","score_opus":0.2972465483711375,"score_gpt":0.4600617848547799,"score_spread":0.16281523648364238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2790826992","genre_codex":"methods","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":"methods","prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09657835,0.001457671,0.8619019,0.023240557,0.0005267592,0.001788734,0.00016205182,0.00039101948,0.013952965],"genre_scores_gemma":[0.46546528,0.0009567801,0.5251057,0.0017966423,0.00023318882,0.004228831,0.00018580242,0.0004096336,0.001618088],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.42463616,0.471051,0.035024084,0.012624402,0.05247451,0.004189831],"domain_scores_gemma":[0.26967245,0.59103554,0.042023376,0.054600567,0.03885159,0.0038164195],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.34820354,0.0017141226,0.0018513553,0.012285952,0.015585089,0.02897031,0.005840728,0.0030087167,0.001998928],"category_scores_gemma":[0.5959751,0.002289175,0.001477392,0.015035531,0.044889778,0.029303452,0.021588143,0.00815405,0.00061717385],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017838133,0.00008108262,0.00860193,0.0013894897,0.00008332356,0.00040551537,0.43077517,0.0021300837,0.0033568605,0.4403273,0.002804291,0.10986666],"study_design_scores_gemma":[0.00010687532,0.0001226668,0.0049673584,0.0043364167,0.00008416948,0.0008140045,0.13052802,0.012599165,0.004692213,0.8022782,0.039070074,0.00040076431],"about_ca_topic_score_codex":0.007475672,"about_ca_topic_score_gemma":0.0061694477,"teacher_disagreement_score":0.65179646,"about_ca_system_score_codex":0.01818968,"about_ca_system_score_gemma":0.04330096,"threshold_uncertainty_score":0.8037811},"labels":[],"label_agreement":null},{"id":"W2791141695","doi":"10.48550/arxiv.1801.02716","title":"The Android Update Problem: An Empirical Study","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Android (operating system); Computer science; Phone; Empirical research; Merge (version control); Android application; Operating system; Information retrieval","score_opus":0.08550174080827891,"score_gpt":0.24521493097975525,"score_spread":0.15971319017147634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2791141695","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9946824,0.0003327892,0.0007151904,0.00065368763,0.000010773571,0.00014789306,0.000382766,0.000026606675,0.003047919],"genre_scores_gemma":[0.9955219,0.00035155434,0.0019772302,0.00023416492,0.000022942791,0.00015027857,0.00095959735,0.00003068071,0.00075174094],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9858526,0.006574161,0.001236579,0.0015940681,0.0037535953,0.0009890061],"domain_scores_gemma":[0.74549943,0.20255448,0.025765672,0.008046625,0.014548624,0.0035851446],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011587248,0.00055262324,0.0005603912,0.0034353388,0.0023804528,0.0028672009,0.0022008973,0.0024308746,0.0034543439],"category_scores_gemma":[0.107265,0.0005815116,0.00054949336,0.0046864087,0.002528124,0.0067401323,0.0023249502,0.0039859503,0.00084505975],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044756613,0.0043384545,0.9086202,0.0007182524,0.00013236239,0.0019039154,0.025416039,0.0013918395,0.00063862524,0.003527414,0.008365106,0.0445002],"study_design_scores_gemma":[0.0001815689,0.0011672361,0.8320097,0.00061518553,0.00016522832,0.0043516587,0.10504846,0.025560107,0.0013032953,0.0028555742,0.026605701,0.0001362426],"about_ca_topic_score_codex":0.0078209415,"about_ca_topic_score_gemma":0.0074354005,"teacher_disagreement_score":0.011587248,"about_ca_system_score_codex":0.0017100306,"about_ca_system_score_gemma":0.0014155734,"threshold_uncertainty_score":0.061279953},"labels":[],"label_agreement":null},{"id":"W2791914593","doi":"10.1002/smr.1936","title":"Merge‐Tree: Visualizing the integration of commits into Linux","year":2018,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Commit; Computer science; Merge (version control); Linux kernel; Operating system; Directed acyclic graph; Android (operating system); Programming language; Database; Parallel computing; Algorithm","score_opus":0.022560551012825312,"score_gpt":0.3253684435478788,"score_spread":0.30280789253505347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2791914593","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25405124,0.0018538992,0.5114467,0.0025525172,0.00068292103,0.0004801896,0.024262201,0.1864565,0.018213782],"genre_scores_gemma":[0.61717266,0.00083168095,0.3493976,0.00026038857,0.000081154045,0.00026206867,0.016961575,0.009956671,0.0050762217],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993788,0.00014288814,0.00005996732,0.000118586824,0.00022957304,0.00007021589],"domain_scores_gemma":[0.99490225,0.002495908,0.0005137994,0.00059299637,0.0009533437,0.0005417638],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015802333,0.0008667236,0.00041111396,0.0046460954,0.00074602256,0.0021332612,0.0011923653,0.0008656844,0.0069445055],"category_scores_gemma":[0.007943914,0.00046292978,0.00066034234,0.0031854725,0.00046002297,0.0028727916,0.0022140257,0.0015067173,0.0014637134],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032418205,0.00084816216,0.08091795,0.0024650146,0.00036466407,0.00318296,0.029077668,0.06085491,0.03664883,0.042550627,0.20881136,0.531036],"study_design_scores_gemma":[0.00032135702,0.000490203,0.05620077,0.00073160784,0.00020163914,0.0016295826,0.006122599,0.62246895,0.036789436,0.045106743,0.22950554,0.00043163792],"about_ca_topic_score_codex":0.010685297,"about_ca_topic_score_gemma":0.010975725,"teacher_disagreement_score":0.010685297,"about_ca_system_score_codex":0.00060298305,"about_ca_system_score_gemma":0.0012197997,"threshold_uncertainty_score":0.023231685},"labels":[],"label_agreement":null},{"id":"W2792248403","doi":"10.11575/prism/25515","title":"Data Analytics for Optimized Matching in Software Development","year":2017,"lang":"en","type":"dissertation","venue":"PRISM (University of Calgary)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates; Alberta Innovates - Technology Futures","keywords":"Computer science; Matching (statistics); Analytics; Data science; Software analytics; Software; Software engineering; Software development; Data mining; Software development process; Statistics; Mathematics; Programming language","score_opus":0.03969259177900502,"score_gpt":0.2738157909532583,"score_spread":0.23412319917425328,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2792248403","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044734627,0.0013344139,0.9314344,0.0025912225,0.00017773097,0.0006614953,0.004634661,0.009195706,0.005235801],"genre_scores_gemma":[0.37959352,0.00079183996,0.6099818,0.0004089903,0.00007752322,0.00072073867,0.0067933383,0.0004881554,0.0011440594],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9885316,0.0040348084,0.0012419086,0.0023648536,0.0033393798,0.00048733983],"domain_scores_gemma":[0.97318584,0.016647886,0.002332471,0.003705736,0.0034711754,0.00065682235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009298523,0.0017130882,0.001668146,0.0067878165,0.0011885371,0.004784102,0.0024075394,0.0014649448,0.003519234],"category_scores_gemma":[0.05594997,0.00096148124,0.0018352877,0.008481373,0.001291428,0.006045055,0.0050330944,0.0027904608,0.0015594696],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008899454,0.000819379,0.03620899,0.0013574715,0.0003276344,0.00057554897,0.0020924243,0.2662051,0.005290563,0.056923985,0.016112622,0.6131964],"study_design_scores_gemma":[0.000053590884,0.00013176726,0.0034734246,0.00015208799,0.000047557605,0.0000862222,0.00079148775,0.88221645,0.004797298,0.095912635,0.012278845,0.000058618927],"about_ca_topic_score_codex":0.007373189,"about_ca_topic_score_gemma":0.00699012,"teacher_disagreement_score":0.009298523,"about_ca_system_score_codex":0.0022944969,"about_ca_system_score_gemma":0.00283328,"threshold_uncertainty_score":0.04917586},"labels":[],"label_agreement":null},{"id":"W2793175118","doi":"10.1007/s10664-018-9604-y","title":"Studying the consistency of star ratings and the complaints in 1 &amp; 2-star user reviews for top free cross-platform Android and iOS apps","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Android (operating system); Computer science; Mobile apps; World Wide Web; Revenue; Star (game theory); App store; Download; User experience design; Consistency (knowledge bases); Human–computer interaction; Operating system; Artificial intelligence","score_opus":0.05466891975716319,"score_gpt":0.32305640697843746,"score_spread":0.26838748722127426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2793175118","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98827994,0.00065651897,0.0032758103,0.0005360902,0.00016306601,0.00020360152,0.0013643061,0.00012777623,0.0053928327],"genre_scores_gemma":[0.9942525,0.00015359167,0.001951582,0.0001938805,0.000085938256,0.00020872675,0.00167971,0.000065222695,0.0014087451],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9356899,0.019501472,0.012108945,0.0053846436,0.025805965,0.0015089944],"domain_scores_gemma":[0.46141055,0.29805914,0.08257044,0.01932023,0.13409822,0.0045414804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04069731,0.00033827455,0.0008852299,0.0071166665,0.0011923113,0.0029582684,0.00097597146,0.0010588277,0.0014785215],"category_scores_gemma":[0.30747947,0.00040413672,0.0011444237,0.00497898,0.0011916097,0.0034557625,0.003042175,0.0011590652,0.00082001224],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010462541,0.00020399087,0.92810506,0.0008216083,0.0007929611,0.00014077472,0.016580047,0.0004384766,0.0030656029,0.00064199726,0.005645695,0.042517498],"study_design_scores_gemma":[0.00002870785,0.00055063434,0.98391336,0.0001299955,0.00019528673,0.0002256902,0.006226279,0.0022095703,0.0014475434,0.00024665872,0.0047473307,0.000078939556],"about_ca_topic_score_codex":0.0027867043,"about_ca_topic_score_gemma":0.005467275,"teacher_disagreement_score":0.04069731,"about_ca_system_score_codex":0.0013588689,"about_ca_system_score_gemma":0.001264213,"threshold_uncertainty_score":0.21523052},"labels":[],"label_agreement":null},{"id":"W2793227253","doi":"10.1007/s10664-017-9592-3","title":"ProMeTA: a taxonomy for program metamodels in program reverse engineering","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Japan Society for the Promotion of Science","keywords":"Metamodeling; Taxonomy (biology); Computer science; Reuse; Software engineering; Program comprehension; Systems engineering; Orthogonality; Artificial intelligence; Engineering; Programming language; Software; Software system","score_opus":0.0615456604311363,"score_gpt":0.3297100329207495,"score_spread":0.2681643724896132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2793227253","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0069016726,0.0036227345,0.9711057,0.0019466772,0.00016043348,0.0013990763,0.001847621,0.0022119528,0.010804081],"genre_scores_gemma":[0.027015118,0.0032058232,0.9613904,0.00045338966,0.000059051134,0.0021253047,0.0033781363,0.0003565154,0.0020161993],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98575145,0.005334822,0.0036301373,0.001623399,0.0030895139,0.0005706397],"domain_scores_gemma":[0.9763856,0.01026515,0.0031494687,0.0046408786,0.0048208954,0.0007379487],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013787038,0.0017147232,0.0011953032,0.017833598,0.0032277298,0.007226504,0.0031681242,0.0030486898,0.0031974595],"category_scores_gemma":[0.026835347,0.0012820405,0.0036632998,0.015955947,0.004181694,0.017243909,0.0053956364,0.005018899,0.0015417153],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000099173085,0.00019467673,0.0068421955,0.0032813877,0.00012693822,0.0005599754,0.0113091115,0.0055861,0.005541444,0.6724969,0.012675701,0.28128648],"study_design_scores_gemma":[0.000051579107,0.00021745895,0.0038297493,0.0054210997,0.00018171176,0.0026397991,0.0051436634,0.022137374,0.003814149,0.36466014,0.5916977,0.00020553442],"about_ca_topic_score_codex":0.008182196,"about_ca_topic_score_gemma":0.009818329,"teacher_disagreement_score":0.017833598,"about_ca_system_score_codex":0.0050313324,"about_ca_system_score_gemma":0.0135206,"threshold_uncertainty_score":0.07291365},"labels":[],"label_agreement":null},{"id":"W2794642134","doi":"10.1145/3178315.3178326","title":"Prioritizing lingering bugs","year":2018,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software bug; Computer science; Security bug; Principal (computer security); Deci-; Risk analysis (engineering); Software; Business; Computer security; Programming language; Political science","score_opus":0.02061314265003952,"score_gpt":0.2630739726971938,"score_spread":0.24246083004715427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2794642134","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64025515,0.00074638904,0.3507166,0.0012221335,0.00011870513,0.00044628346,0.0002178571,0.0021868818,0.004090058],"genre_scores_gemma":[0.9646176,0.00010671009,0.033478305,0.00010399975,0.00003405336,0.000052088344,0.00011435097,0.000053290747,0.0014395523],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99688303,0.00089275755,0.00022936473,0.00067816227,0.000940978,0.00037572125],"domain_scores_gemma":[0.96808136,0.020554183,0.0052741463,0.0015687554,0.0029805168,0.0015410766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005235185,0.0012721724,0.0011287616,0.0017452636,0.00049727556,0.0015672789,0.0015479584,0.0014808134,0.0024464268],"category_scores_gemma":[0.034094363,0.00066113827,0.0005951619,0.00061620586,0.0007174282,0.0018967052,0.0009977387,0.0016847701,0.0003194968],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077876984,0.0011794586,0.10600468,0.00045333922,0.00025089193,0.0005772498,0.0004834943,0.5787098,0.01777171,0.008455889,0.0024472324,0.28288758],"study_design_scores_gemma":[0.00004082305,0.0005273309,0.00951849,0.000032393382,0.000059237835,0.00017716516,0.00007795463,0.97796905,0.0037208311,0.007291696,0.0005484368,0.000036557514],"about_ca_topic_score_codex":0.004731305,"about_ca_topic_score_gemma":0.006483986,"teacher_disagreement_score":0.005235185,"about_ca_system_score_codex":0.001281516,"about_ca_system_score_gemma":0.002323548,"threshold_uncertainty_score":0.027686596},"labels":[],"label_agreement":null},{"id":"W2794757881","doi":"10.1145/3178315.3178324","title":"Stakeholder Concern-Driven Requirements Analytics","year":2018,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Requirements management; Computer science; Requirements analysis; Requirements engineering; Analytics; Traceability; Business requirements; Process (computing); Stakeholder; Requirements elicitation; Process management; Requirement; Risk analysis (engineering); Business process; Software engineering; Engineering; Data science; Software; Work in process; Business","score_opus":0.09397902749149702,"score_gpt":0.30309292860879433,"score_spread":0.20911390111729733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2794757881","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046996914,0.00027779568,0.9157761,0.0035482338,0.00005973668,0.0009394803,0.0016368896,0.0032886853,0.027476206],"genre_scores_gemma":[0.48559546,0.0003919363,0.50381285,0.00046246612,0.00005062836,0.0007175214,0.0038322427,0.00047324493,0.0046636453],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9855525,0.007849258,0.0007140182,0.0009376511,0.0045027346,0.0004438711],"domain_scores_gemma":[0.9818897,0.008212723,0.0015084785,0.0024143714,0.0054898304,0.00048501737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009997496,0.0012854371,0.00056461495,0.003940199,0.00081452064,0.0042320397,0.0017887735,0.0010592589,0.0033507068],"category_scores_gemma":[0.028757839,0.0005203846,0.0009542865,0.0030739796,0.0007564678,0.0048124464,0.0039024875,0.0014754933,0.0012006137],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000427981,0.00066691096,0.0243458,0.0017783842,0.00022038621,0.0011957643,0.020662688,0.05430639,0.037645046,0.28916898,0.03419325,0.53538847],"study_design_scores_gemma":[0.00006542986,0.0002999538,0.01142406,0.00057731546,0.00008446034,0.0009268085,0.0155796865,0.50874484,0.033142928,0.26445588,0.16447765,0.00022090416],"about_ca_topic_score_codex":0.0017915709,"about_ca_topic_score_gemma":0.002184304,"teacher_disagreement_score":0.009997496,"about_ca_system_score_codex":0.0016904894,"about_ca_system_score_gemma":0.0021489535,"threshold_uncertainty_score":0.05287248},"labels":[],"label_agreement":null},{"id":"W2795368212","doi":"10.1109/saner.2018.8330213","title":"Classifying stack overflow posts on API issues","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Queen's University","funders":"","keywords":"Computer science; Documentation; Task (project management); Conditional random field; Generalizability theory; Field (mathematics); World Wide Web; Application programming interface; Baseline (sea); Data science; Software engineering; Artificial intelligence; Programming language; Engineering","score_opus":0.03463733268093194,"score_gpt":0.3190475714169406,"score_spread":0.2844102387360087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795368212","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94891864,0.00091343716,0.03403548,0.000587213,0.00023815913,0.00043033963,0.005610092,0.003334957,0.00593162],"genre_scores_gemma":[0.93504286,0.00048562788,0.048021447,0.00023079317,0.000293799,0.000285436,0.010845492,0.00028200055,0.004512518],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99876213,0.0002591073,0.00016412225,0.00024290374,0.0004403501,0.00013137746],"domain_scores_gemma":[0.98504907,0.008540434,0.0027844396,0.0006596298,0.0025140794,0.00045228685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001306859,0.00062108523,0.00044284112,0.0040630233,0.0007460322,0.00097406656,0.00048986694,0.0007870976,0.001485066],"category_scores_gemma":[0.012491777,0.00019223061,0.0004672806,0.0019480601,0.00039390964,0.00225551,0.0008087523,0.0009564816,0.0009338421],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019659845,0.0008441123,0.28025368,0.002427984,0.00018372716,0.0023041153,0.010171613,0.0046284674,0.09255518,0.0034555704,0.04419259,0.557017],"study_design_scores_gemma":[0.0000803663,0.0013371473,0.6113521,0.0005060816,0.00043332684,0.002218022,0.00886618,0.24300395,0.058545604,0.0048620757,0.06853265,0.00026252065],"about_ca_topic_score_codex":0.0031914937,"about_ca_topic_score_gemma":0.0074404064,"teacher_disagreement_score":0.0040630233,"about_ca_system_score_codex":0.0005704734,"about_ca_system_score_gemma":0.0009851411,"threshold_uncertainty_score":0.0069114566},"labels":[],"label_agreement":null},{"id":"W2795374451","doi":"10.1109/saner.2018.8330193","title":"Design patterns impact on software quality: Where are the theories?","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Polytechnique Montréal","funders":"","keywords":"Software engineering; Computer science; Software development; Software quality; Software; Software peer review; Business process reengineering; Software design; Quality (philosophy); Software construction; Engineering; Programming language","score_opus":0.043395890520943274,"score_gpt":0.3347946077419956,"score_spread":0.2913987172210523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795374451","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5524284,0.07247672,0.117955975,0.14032091,0.0010821851,0.0004365883,0.0008894107,0.000860012,0.11354983],"genre_scores_gemma":[0.96512604,0.011735611,0.0156941,0.0035906478,0.00033831142,0.00018231987,0.00028489408,0.00027668767,0.0027714258],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9705771,0.012072957,0.0019133257,0.0027810782,0.011276173,0.0013792291],"domain_scores_gemma":[0.7925448,0.13576329,0.027182342,0.0145949675,0.02614003,0.0037745736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015245392,0.0013474283,0.0011126646,0.004390758,0.0016374718,0.007429758,0.001938574,0.0028986298,0.0070393058],"category_scores_gemma":[0.11524695,0.0009449333,0.0018724494,0.0073125325,0.0053136293,0.014626132,0.004115214,0.004190408,0.001144463],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051223894,0.0012380764,0.282078,0.002951439,0.0011631967,0.00024667982,0.013804329,0.008061459,0.0016385206,0.20575835,0.01291938,0.46962836],"study_design_scores_gemma":[0.00038241368,0.0010982536,0.26085377,0.0047998326,0.0016527277,0.00039300072,0.011863968,0.022940496,0.0043212464,0.62511814,0.06636637,0.00020966977],"about_ca_topic_score_codex":0.0080054095,"about_ca_topic_score_gemma":0.0046331915,"teacher_disagreement_score":0.015245392,"about_ca_system_score_codex":0.0061297924,"about_ca_system_score_gemma":0.003776689,"threshold_uncertainty_score":0.08062631},"labels":[],"label_agreement":null},{"id":"W2795591975","doi":"10.1109/saner.2018.8330230","title":"Detection of protection-impacting changes during software evolution","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Software evolution; Computer science; Software; Software system; Software construction; Operating system","score_opus":0.01986508054890323,"score_gpt":0.2510046316356482,"score_spread":0.23113955108674494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795591975","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93987197,0.00041198867,0.055035666,0.000062562576,0.00002400865,0.00012730673,0.0007121962,0.0027962846,0.0009580642],"genre_scores_gemma":[0.9526815,0.00014451238,0.045107678,0.000021536336,0.00001335259,0.00006041248,0.001407222,0.000107408356,0.0004563599],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99613404,0.0005035273,0.0003102751,0.00096751156,0.001883252,0.00020137469],"domain_scores_gemma":[0.98324513,0.007141477,0.004411855,0.002273155,0.0025484674,0.00038000254],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015533328,0.00059154426,0.0005121417,0.0041116336,0.00034031758,0.0007094652,0.0006435586,0.0007098885,0.0005430308],"category_scores_gemma":[0.014018679,0.00033950582,0.00058822014,0.001901616,0.00033085915,0.00090698904,0.00075495074,0.0004964264,0.0002558959],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064409786,0.00042773676,0.5985037,0.00048472598,0.0002159405,0.0020616911,0.0012942685,0.019253148,0.083261594,0.0007421097,0.0013534073,0.29175764],"study_design_scores_gemma":[0.00004256919,0.00083319534,0.5992677,0.00008177592,0.00028222555,0.0030988182,0.00043885334,0.31745645,0.0736089,0.001131829,0.0036748806,0.00008286661],"about_ca_topic_score_codex":0.0028175078,"about_ca_topic_score_gemma":0.0032596597,"teacher_disagreement_score":0.0041116336,"about_ca_system_score_codex":0.00040156124,"about_ca_system_score_gemma":0.00066585053,"threshold_uncertainty_score":0.008214891},"labels":[],"label_agreement":null},{"id":"W2795698999","doi":"10.1109/saner.2018.8330233","title":"Towards just-in-time suggestions for log changes (journal-first abstract)","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Queen's University","funders":"","keywords":"Commit; Computer science; Code (set theory); Open source; Source code; Software; Data science; Software engineering; Database; Programming language; Set (abstract data type)","score_opus":0.03734781818199455,"score_gpt":0.30597578561755834,"score_spread":0.2686279674355638,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795698999","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2428453,0.0052265595,0.6041842,0.021264559,0.007527544,0.0019468932,0.0073073665,0.084788814,0.024908697],"genre_scores_gemma":[0.5617097,0.0011212465,0.41775373,0.002195232,0.0014985802,0.00053759915,0.0050710193,0.0027828263,0.0073301066],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.95908695,0.020960318,0.0031538112,0.004931488,0.011052521,0.0008148619],"domain_scores_gemma":[0.45975202,0.33929732,0.048086017,0.057029698,0.08992517,0.005909798],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031909548,0.0013068365,0.0010447198,0.0075355456,0.0014315349,0.011187612,0.002553238,0.003941977,0.010928536],"category_scores_gemma":[0.3965581,0.0011472529,0.0007562134,0.005099772,0.0009991704,0.009784864,0.0031412467,0.0032495353,0.013011491],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008843068,0.0011411227,0.09128035,0.001488959,0.00020947274,0.0004775484,0.0026820553,0.0052276254,0.014389927,0.003138878,0.093234636,0.78584516],"study_design_scores_gemma":[0.0011359305,0.0024011224,0.17223331,0.0025343816,0.00080157386,0.00268263,0.007003176,0.37173858,0.08678548,0.040119153,0.31125084,0.0013138226],"about_ca_topic_score_codex":0.0027745927,"about_ca_topic_score_gemma":0.0038937703,"teacher_disagreement_score":0.031909548,"about_ca_system_score_codex":0.0011441475,"about_ca_system_score_gemma":0.004306577,"threshold_uncertainty_score":0.16875577},"labels":[],"label_agreement":null},{"id":"W2795734224","doi":"10.1145/3173574.3173859","title":"Leveraging Community-Generated Videos and Command Logs to Classify and Recommend Software Workflows","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Autodesk (Canada)","funders":"","keywords":"Workflow; Computer science; Task (project management); Software; Recommender system; Data science; Software engineering; Data mining; World Wide Web; Database; Engineering","score_opus":0.04597109284609218,"score_gpt":0.28601824589909614,"score_spread":0.24004715305300395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795734224","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5566501,0.0012522076,0.4288843,0.0007277328,0.00014709195,0.0009958944,0.0026944366,0.0043021594,0.004346055],"genre_scores_gemma":[0.7643746,0.00034760547,0.22692919,0.00009088476,0.00009509691,0.0002732966,0.005485395,0.00017043995,0.002233429],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986357,0.0004655241,0.00007994472,0.00042223214,0.00029362476,0.000102959886],"domain_scores_gemma":[0.9874791,0.0071229534,0.0009083844,0.0011011964,0.0028095734,0.00057887635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031581016,0.0009650614,0.00064419163,0.004774212,0.00072237843,0.0011145731,0.0011461792,0.00087934727,0.000859246],"category_scores_gemma":[0.018205402,0.0002940687,0.000653795,0.001980237,0.00030863687,0.00249771,0.00079486327,0.0010080381,0.0005992348],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013054735,0.0018046323,0.12719585,0.001027332,0.00045160274,0.00048608225,0.0023125068,0.05870665,0.026388427,0.0031636355,0.013887313,0.76327044],"study_design_scores_gemma":[0.00008478482,0.0006397588,0.033039294,0.00011167842,0.00016101504,0.0002761114,0.001138813,0.9425008,0.009838768,0.004524825,0.0075868,0.00009732046],"about_ca_topic_score_codex":0.018671546,"about_ca_topic_score_gemma":0.039726496,"teacher_disagreement_score":0.018671546,"about_ca_system_score_codex":0.0007686675,"about_ca_system_score_gemma":0.0013215975,"threshold_uncertainty_score":0.037125707},"labels":[],"label_agreement":null},{"id":"W2795753518","doi":"10.1109/saner.2018.8330219","title":"Syntax and sensibility: Using language models to detect and correct syntax errors","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs; Nvidia","keywords":"Computer science; Syntax error; Syntax; Programming language; Parsing; Natural language processing; Java; Abstract syntax; Abstract syntax tree; Artificial intelligence; Language model; Source code; Security token; Code (set theory)","score_opus":0.035670842053011735,"score_gpt":0.29222504636629054,"score_spread":0.2565542043132788,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795753518","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15619443,0.00117533,0.7470572,0.0019475905,0.00052863074,0.00037997632,0.0028839617,0.08408834,0.0057444526],"genre_scores_gemma":[0.5621796,0.00048036056,0.4238623,0.0008803104,0.000110300054,0.00021233589,0.003750717,0.0036023522,0.0049216016],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969398,0.0012852096,0.00021428535,0.0009285629,0.00045577082,0.00017637596],"domain_scores_gemma":[0.9880938,0.0070476667,0.001173675,0.0018693203,0.0015096092,0.0003059712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037359046,0.0028231323,0.00087931374,0.0024998114,0.00050755043,0.002518144,0.002372957,0.002068344,0.0040025376],"category_scores_gemma":[0.028396703,0.0009937031,0.0014343228,0.0010224393,0.0011951278,0.0072487774,0.0024196543,0.0030366902,0.0035843686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008891474,0.0004830768,0.029856471,0.0011010679,0.0004040132,0.000934588,0.0026922412,0.08795638,0.04403979,0.004267092,0.020438535,0.80693763],"study_design_scores_gemma":[0.0001159383,0.0004827396,0.008162111,0.00027983807,0.00026608814,0.0006835492,0.000752332,0.9142117,0.040031586,0.022987276,0.011814214,0.00021261285],"about_ca_topic_score_codex":0.0066210357,"about_ca_topic_score_gemma":0.008812121,"teacher_disagreement_score":0.0066210357,"about_ca_system_score_codex":0.0010255118,"about_ca_system_score_gemma":0.0022080473,"threshold_uncertainty_score":0.019757569},"labels":[],"label_agreement":null},{"id":"W2795757906","doi":"10.1109/saner.2018.8330196","title":"Micro-clones in evolving software","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; clone (Java method); Software maintenance; Code (set theory); Computer science; BitTorrent tracker; Base (topology); Software; Programming language; Source code; Source lines of code; Software system; Biology; Artificial intelligence; Genetics; Gene; Mathematics","score_opus":0.015377828643177674,"score_gpt":0.26542894900770336,"score_spread":0.2500511203645257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795757906","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88648105,0.0024046535,0.107828505,0.0001850466,0.00003983864,0.00014603534,0.00025450488,0.0009858377,0.0016744415],"genre_scores_gemma":[0.9547343,0.00053728314,0.04309051,0.00008513084,0.000022260803,0.00007573427,0.00037123763,0.0001372755,0.00094620825],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99403137,0.0011349015,0.00044679857,0.0015489869,0.002580408,0.00025749343],"domain_scores_gemma":[0.94551224,0.02898009,0.012391797,0.005395111,0.0068792617,0.00084156916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002932359,0.00038774105,0.0004919204,0.0041159764,0.00065534277,0.0015420918,0.00072478846,0.00072322704,0.00060246635],"category_scores_gemma":[0.03626968,0.0004575867,0.0004264887,0.0029424964,0.0011665857,0.0025860057,0.0014156495,0.0007102526,0.00019651589],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035290423,0.00010399824,0.504233,0.0005560974,0.0002143552,0.0021321874,0.0061938753,0.02213303,0.03322988,0.008816134,0.0016120255,0.42042243],"study_design_scores_gemma":[0.00004929859,0.0007420582,0.6440637,0.00035148385,0.00049657107,0.009983294,0.004009815,0.21740988,0.06584785,0.03282916,0.024020553,0.0001962696],"about_ca_topic_score_codex":0.0032132294,"about_ca_topic_score_gemma":0.004130889,"teacher_disagreement_score":0.0041159764,"about_ca_system_score_codex":0.00080337253,"about_ca_system_score_gemma":0.00075843907,"threshold_uncertainty_score":0.0155079365},"labels":[],"label_agreement":null},{"id":"W2795868997","doi":"10.1109/saner.2018.8330192","title":"Ten years of JDeodorant: Lessons learned from the hunt for smells","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; Code smell; Computer science; Java; Code (set theory); Software engineering; Source code; Software; Empirical research; Programming language; Software development; Software quality","score_opus":0.07313347573389459,"score_gpt":0.32823969725265145,"score_spread":0.25510622151875684,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2795868997","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44119996,0.14190876,0.19867541,0.12115047,0.0052743056,0.00045362892,0.009263743,0.01360133,0.06847244],"genre_scores_gemma":[0.49286398,0.07386144,0.36460415,0.01555715,0.0018403828,0.00028410694,0.016210018,0.008588595,0.0261902],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9928042,0.0015945424,0.00062207266,0.0015569231,0.0030255348,0.00039672074],"domain_scores_gemma":[0.9223098,0.04832745,0.0032515123,0.012417888,0.010941079,0.0027523132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013901713,0.0008962841,0.0009622959,0.0055271704,0.0015222722,0.0069699017,0.0028878683,0.0017772991,0.0039511938],"category_scores_gemma":[0.06205356,0.0009772503,0.0011829975,0.0041252887,0.0034222875,0.019256188,0.004393302,0.0050231325,0.001915684],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036546926,0.00037518804,0.041782446,0.0016568437,0.0001221836,0.00042994416,0.006890115,0.0014384848,0.0037350794,0.012099304,0.066075325,0.8650296],"study_design_scores_gemma":[0.000094440016,0.00054949196,0.060234126,0.0030168267,0.00017203669,0.0018426022,0.010908138,0.008092518,0.013501126,0.04403049,0.8571787,0.00037955408],"about_ca_topic_score_codex":0.0070238444,"about_ca_topic_score_gemma":0.016818069,"teacher_disagreement_score":0.013901713,"about_ca_system_score_codex":0.0019872596,"about_ca_system_score_gemma":0.0025744776,"threshold_uncertainty_score":0.073520124},"labels":[],"label_agreement":null},{"id":"W2796071471","doi":"10.1109/saner.2018.8330217","title":"A generalized model for visualizing library popularity, adoption, and diffusion within a software ecosystem","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science","keywords":"Computer science; Popularity; Dependency (UML); Software; Reuse; Ecosystem; Data science; Software evolution; Software engineering; World Wide Web; Software development; Software construction; Ecology; Programming language","score_opus":0.03195126229654964,"score_gpt":0.28509201897984204,"score_spread":0.2531407566832924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2796071471","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47980714,0.00048233874,0.5048169,0.0015253782,0.00003640462,0.00015084649,0.0028498236,0.00341786,0.006913373],"genre_scores_gemma":[0.90740144,0.00022473394,0.08887766,0.00006628838,0.000010690244,0.00021792973,0.0011773914,0.00014328046,0.0018805118],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999353,0.0003059504,0.000036168178,0.00015960325,0.000079060585,0.00006613116],"domain_scores_gemma":[0.9966336,0.0020470594,0.00044843263,0.00041671432,0.00028852627,0.00016570628],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0014569868,0.0005672035,0.0004218209,0.0037540758,0.0005440857,0.0023622585,0.0010957406,0.0011702167,0.0024743408],"category_scores_gemma":[0.008326993,0.00040749428,0.0010370368,0.0037937148,0.000923123,0.0035306262,0.0011149623,0.0008526568,0.00049678166],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038521414,0.0001894966,0.111653455,0.00031611725,0.00025927878,0.0005097319,0.0045722392,0.639338,0.0071210614,0.14211316,0.006769968,0.0867722],"study_design_scores_gemma":[0.000016343658,0.000055855955,0.009104497,0.000018732808,0.000026592943,0.00014354556,0.0004606882,0.95091254,0.00044975936,0.03574876,0.0030355356,0.000027167865],"about_ca_topic_score_codex":0.02624758,"about_ca_topic_score_gemma":0.026599275,"teacher_disagreement_score":0.9962459,"about_ca_system_score_codex":0.0017012552,"about_ca_system_score_gemma":0.000819588,"threshold_uncertainty_score":0.05218965},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"simulation_or_modeling","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"simulation_or_modeling","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"agree"},{"id":"W2796141603","doi":"10.1109/saner.2018.8330194","title":"Benchmarks for software clone detection: A ten-year retrospective","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Saskatchewan","funders":"","keywords":"clone (Java method); Benchmark (surveying); Java; Computer science; Source lines of code; Software maintenance; Cloning (programming); Software; Software engineering; Precision and recall; Code refactoring; Set (abstract data type); Empirical research; Source code; Code (set theory); Software system; Programming language; Machine learning; Statistics; Biology","score_opus":0.011739618515613823,"score_gpt":0.2597864749169241,"score_spread":0.24804685640131025,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2796141603","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32580003,0.25800225,0.1302719,0.016612293,0.0067339772,0.0023945125,0.17693122,0.030621333,0.052632466],"genre_scores_gemma":[0.32206523,0.04793841,0.1367048,0.005203329,0.0011793458,0.0031592217,0.4693242,0.006347252,0.0080781905],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.93603146,0.012410551,0.007894308,0.009686862,0.03171655,0.0022602824],"domain_scores_gemma":[0.7111301,0.08232217,0.02448241,0.034460135,0.13938767,0.008217472],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04421092,0.0027359491,0.0023529194,0.026276741,0.0024795942,0.0074294102,0.006297709,0.0021925997,0.0023218058],"category_scores_gemma":[0.18220694,0.0015030339,0.0023706723,0.025769657,0.0024722714,0.009177858,0.006312492,0.0045714807,0.0031607929],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008263298,0.0007075453,0.11131987,0.0062781093,0.0008191076,0.00026603398,0.0013504452,0.0073566297,0.004014965,0.0123422565,0.30222455,0.5524942],"study_design_scores_gemma":[0.00022538783,0.0017253229,0.1916156,0.00808041,0.00088218163,0.0020427203,0.0021638894,0.024505029,0.022999724,0.012313365,0.73292226,0.00052418537],"about_ca_topic_score_codex":0.013832014,"about_ca_topic_score_gemma":0.016846122,"teacher_disagreement_score":0.04421092,"about_ca_system_score_codex":0.0070497124,"about_ca_system_score_gemma":0.005641703,"threshold_uncertainty_score":0.23381251},"labels":[],"label_agreement":null},{"id":"W2796224197","doi":"10.1109/saner.2018.8330237","title":"The relationship between evolutionary coupling and defects in large industrial software (journal-first abstract)","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software; Coupling (piping); Computer science; Software evolution; Software engineering; Software system; Engineering; Software construction; Operating system; Mechanical engineering","score_opus":0.06366256133284046,"score_gpt":0.296023197120436,"score_spread":0.23236063578759555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2796224197","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99860954,0.000098875666,0.00075932493,0.00003257223,0.0000017342101,0.0000040087707,0.000017510494,0.000012316835,0.00046397702],"genre_scores_gemma":[0.9995766,0.000022059809,0.00028615864,0.0000078247795,0.0000030937347,0.0000019630404,0.000021968253,0.0000055276737,0.00007469316],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9976948,0.0008882345,0.00021862469,0.00030522377,0.0006843757,0.00020877634],"domain_scores_gemma":[0.83058393,0.111258164,0.040585257,0.0059485803,0.0074165193,0.0042075645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002619853,0.00029900196,0.00022885423,0.0023843143,0.0002939123,0.0010436185,0.0006832543,0.0006973895,0.0017241875],"category_scores_gemma":[0.057465017,0.00020706234,0.00026771755,0.0014466151,0.0007852924,0.00117019,0.0011081421,0.0009917638,0.00018054561],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003139227,0.0002282998,0.9745318,0.000051708535,0.00011907087,0.00026306862,0.00034775113,0.004141682,0.004350728,0.0005878638,0.00007613046,0.014988082],"study_design_scores_gemma":[0.000009105468,0.00029716987,0.9930565,0.000009751866,0.00003881783,0.00038310385,0.0001397832,0.004253962,0.0011232222,0.0005813592,0.00009826706,0.000009010765],"about_ca_topic_score_codex":0.0017438784,"about_ca_topic_score_gemma":0.0015647068,"teacher_disagreement_score":0.002619853,"about_ca_system_score_codex":0.0004505292,"about_ca_system_score_gemma":0.00035114557,"threshold_uncertainty_score":0.013855219},"labels":[],"label_agreement":null},{"id":"W2796461266","doi":"10.1016/j.jss.2018.03.053","title":"Characterizing and predicting blocking bugs in open source projects","year":2018,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Concordia University","funders":"","keywords":"Blocking (statistics); Software bug; Computer science; Software; Open source; Code (set theory); Source lines of code; Operating system; Programming language; Computer network","score_opus":0.028943378719903365,"score_gpt":0.2765328254276833,"score_spread":0.24758944670777994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2796461266","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9929427,0.00036403068,0.00509082,0.00010526732,0.000020359785,0.000027829057,0.00032129636,0.00054468674,0.000583027],"genre_scores_gemma":[0.9914445,0.00014006988,0.0070450315,0.000030591844,0.000018177398,0.000023776334,0.0008218343,0.00008447549,0.00039144233],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9942275,0.0009814366,0.0008335884,0.0011019792,0.0022503422,0.00060517393],"domain_scores_gemma":[0.85399324,0.083355315,0.0380932,0.006812878,0.013397286,0.004348129],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046381303,0.0006976517,0.00045289786,0.0065053445,0.0007210806,0.0014954204,0.0012029152,0.0015050109,0.0008704897],"category_scores_gemma":[0.070788264,0.0007603383,0.0006158559,0.002930777,0.00082304096,0.0027058676,0.0014318492,0.0012277091,0.00038392516],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002552969,0.000264368,0.94753736,0.00013732888,0.000095963675,0.00021942226,0.00045283127,0.005451977,0.004048738,0.0007676029,0.0011023426,0.039666872],"study_design_scores_gemma":[0.00008355399,0.0010096665,0.74609673,0.00019949391,0.00043867217,0.0016883739,0.0014550613,0.22706614,0.009413183,0.009078528,0.0033731076,0.00009746333],"about_ca_topic_score_codex":0.0077175703,"about_ca_topic_score_gemma":0.012361258,"teacher_disagreement_score":0.0077175703,"about_ca_system_score_codex":0.0007430441,"about_ca_system_score_gemma":0.0020618693,"threshold_uncertainty_score":0.0245291},"labels":[],"label_agreement":null},{"id":"W2800618921","doi":"10.22215/etd/2006-08010","title":"Software product architectural integrity during organizational distress","year":2006,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Canadian Heritage","funders":"","keywords":"Product (mathematics); Computer science; Engineering; Software engineering; Mathematics","score_opus":0.0077561612484027525,"score_gpt":0.24650654592054655,"score_spread":0.2387503846721438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2800618921","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9146542,0.00044882248,0.057993777,0.0024150957,0.000080118574,0.00011373374,0.00023957733,0.0011304142,0.022924313],"genre_scores_gemma":[0.98577833,0.00010629973,0.009408133,0.000089775815,0.000013545514,0.000023421206,0.0002645138,0.00021745326,0.0040985597],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9883183,0.0036132287,0.0009358703,0.0010592752,0.0049088225,0.0011645083],"domain_scores_gemma":[0.93587804,0.016280964,0.015076822,0.015908943,0.014747653,0.002107664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008999141,0.00034093778,0.00042645456,0.0016351779,0.0019529513,0.0043098396,0.0012679,0.0011504836,0.0030714646],"category_scores_gemma":[0.0666734,0.00064155273,0.00028248242,0.0020196412,0.002424567,0.0058800564,0.0047658654,0.0017037408,0.0007946308],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016930335,0.00045130876,0.1849139,0.0005188705,0.0001349156,0.0029804506,0.077652976,0.030442724,0.03471728,0.09563385,0.011852323,0.5590084],"study_design_scores_gemma":[0.00012668392,0.0016836164,0.32472017,0.0007233368,0.00040218295,0.0039880397,0.10182977,0.21964967,0.067629054,0.13149516,0.14751995,0.00023231079],"about_ca_topic_score_codex":0.009570937,"about_ca_topic_score_gemma":0.009771166,"teacher_disagreement_score":0.009570937,"about_ca_system_score_codex":0.0022387165,"about_ca_system_score_gemma":0.0045472533,"threshold_uncertainty_score":0.04759258},"labels":[],"label_agreement":null},{"id":"W2800786079","doi":"","title":"A FrameNet-based Approach for Annotating Natural Language Descriptions of Software Requirements","year":2018,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"FrameNet; Computer science; Natural language; Natural language processing; Natural (archaeology); Software; Programming language; Artificial intelligence; Linguistics; Parsing; History","score_opus":0.09684838032586135,"score_gpt":0.32177982003663014,"score_spread":0.2249314397107688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2800786079","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011651766,0.00031493566,0.9658825,0.0007751255,0.0002502216,0.0010918332,0.0063237064,0.0027039398,0.011006044],"genre_scores_gemma":[0.056135476,0.00042259757,0.92176217,0.00025517936,0.00007355027,0.0018903906,0.0130367875,0.00063276,0.005791147],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9941095,0.0027080313,0.00080368156,0.0009735585,0.0011693499,0.00023590525],"domain_scores_gemma":[0.99124557,0.0040021227,0.00083330175,0.0012132147,0.0024508527,0.00025486253],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006374147,0.0014421914,0.00084202545,0.011421994,0.003204251,0.0029403737,0.0017406688,0.0021227873,0.0064366506],"category_scores_gemma":[0.012780912,0.00097312505,0.0015744834,0.0076171076,0.0023949519,0.00586896,0.0037904617,0.0025132298,0.0022698478],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006646089,0.00046113034,0.004090164,0.0033469764,0.00020108823,0.0028912239,0.028558265,0.021427235,0.05689682,0.45266873,0.055285193,0.3735086],"study_design_scores_gemma":[0.000107034866,0.00022748207,0.00528831,0.0014468103,0.00018835902,0.001647309,0.009313436,0.13792898,0.03227088,0.20806395,0.6032058,0.000311676],"about_ca_topic_score_codex":0.019901639,"about_ca_topic_score_gemma":0.036076937,"teacher_disagreement_score":0.019901639,"about_ca_system_score_codex":0.0038528813,"about_ca_system_score_gemma":0.0048164353,"threshold_uncertainty_score":0.039571583},"labels":[],"label_agreement":null},{"id":"W2801123178","doi":"","title":"The Impact of Operating Systems and Environments on Build Results","year":2017,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Humanities; Political science; Philosophy","score_opus":0.014433439026611134,"score_gpt":0.2682091733719221,"score_spread":0.253775734345311,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2801123178","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9828749,0.0009637038,0.0075960816,0.00036532298,0.000041975796,0.000057370497,0.0018595675,0.0010143893,0.0052267457],"genre_scores_gemma":[0.9909185,0.00029431764,0.0053495327,0.00006388082,0.000027744529,0.00004075355,0.0020070586,0.00052914635,0.00076895504],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9691705,0.0064838384,0.002058669,0.0034069652,0.016884997,0.0019949537],"domain_scores_gemma":[0.69284683,0.22055076,0.034217466,0.022103395,0.02726462,0.0030169326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017669689,0.00097320386,0.000669051,0.005751896,0.00087927683,0.004570264,0.0013164324,0.0008012975,0.0016601479],"category_scores_gemma":[0.12785928,0.0007474597,0.0008521106,0.0055898908,0.0020181185,0.0049728327,0.0026220407,0.0019205906,0.0008039116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010637486,0.00028630678,0.79000944,0.000731117,0.0004877377,0.0007236336,0.0035786405,0.041149843,0.008470707,0.0032067045,0.0043333424,0.14595884],"study_design_scores_gemma":[0.000018718027,0.00052627467,0.94392407,0.00014843792,0.00019033266,0.00064251246,0.002821442,0.029346695,0.011796282,0.0028454128,0.007602094,0.00013774452],"about_ca_topic_score_codex":0.0046477807,"about_ca_topic_score_gemma":0.004741114,"teacher_disagreement_score":0.017669689,"about_ca_system_score_codex":0.0015833706,"about_ca_system_score_gemma":0.0012691596,"threshold_uncertainty_score":0.09344739},"labels":[],"label_agreement":null},{"id":"W2803395207","doi":"10.1109/tse.2018.2838131","title":"Use and Misuse of Continuous Integration Features: An Empirical Study of Projects That (Mis)Use Travis CI","year":2018,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Software deployment; Feature (linguistics); Software; Source code; Code (set theory); Software engineering; Programming language","score_opus":0.05555397349595297,"score_gpt":0.30736840377452834,"score_spread":0.25181443027857536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2803395207","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9991161,0.0000612263,0.00041488896,0.000035690035,0.0000014112828,0.000014614204,0.00008840962,0.000016064345,0.0002515536],"genre_scores_gemma":[0.99802256,0.00009765325,0.0009862912,0.000027720795,0.00000468124,0.0000485651,0.0003577196,0.000025392124,0.0004295006],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9909128,0.0027804498,0.0012131592,0.0014588068,0.0030842384,0.0005504718],"domain_scores_gemma":[0.8561314,0.08041044,0.038211107,0.0073974603,0.014613304,0.0032363513],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0077846786,0.00035831126,0.0003362513,0.0031222622,0.00082943344,0.0013290169,0.00078760664,0.0007772349,0.0008021363],"category_scores_gemma":[0.056530483,0.0004231651,0.00026147984,0.0028844664,0.0013932197,0.0025252106,0.0017722,0.0010714469,0.0003842243],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015508357,0.00030703796,0.9522304,0.00014542365,0.000045104567,0.0005428843,0.01488865,0.00029961552,0.0019719317,0.0002145497,0.00048655967,0.028712794],"study_design_scores_gemma":[0.0000120884,0.00057964836,0.9783407,0.00007201128,0.000029245939,0.0011355433,0.01279694,0.002370702,0.0019921977,0.00020356308,0.0024314648,0.00003599395],"about_ca_topic_score_codex":0.002271158,"about_ca_topic_score_gemma":0.0033834856,"teacher_disagreement_score":0.99221534,"about_ca_system_score_codex":0.00064711424,"about_ca_system_score_gemma":0.0006972353,"threshold_uncertainty_score":0.041169822},"labels":[],"label_agreement":null},{"id":"W2803893519","doi":"10.1109/tse.2018.2836450","title":"Automatically Categorizing Software Technologies","year":2018,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Categorization; Software; Taxonomy (biology); Domain (mathematical analysis); Phrase; Software engineering; Programming language; Natural language processing; Artificial intelligence","score_opus":0.013858434838389265,"score_gpt":0.24211322902831106,"score_spread":0.22825479418992178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2803893519","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2261531,0.0029250653,0.724333,0.0006666827,0.00038480788,0.001414602,0.014435169,0.013715809,0.015971793],"genre_scores_gemma":[0.27719933,0.0012755162,0.68350905,0.00019971713,0.00010827297,0.0014377659,0.030818602,0.0010296704,0.004422052],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9924078,0.0016468347,0.0011698277,0.0014026036,0.0029420375,0.00043100023],"domain_scores_gemma":[0.9766027,0.011469292,0.0024795195,0.002776133,0.006188128,0.00048426844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040957723,0.0013927474,0.001135503,0.03197317,0.0016499653,0.003374484,0.0017287715,0.0018544736,0.0026543115],"category_scores_gemma":[0.031396776,0.0005642247,0.0013240745,0.011563298,0.00077714375,0.008226857,0.0045831418,0.0012157346,0.0020838412],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030492365,0.0002298173,0.040111065,0.0037219222,0.00022142073,0.0011402427,0.007491906,0.004747843,0.050760537,0.035984628,0.021346934,0.8339387],"study_design_scores_gemma":[0.0001424075,0.00040912887,0.061425675,0.0026065477,0.000419847,0.0045920126,0.0146947075,0.32283607,0.10658189,0.13491079,0.3509709,0.0004100663],"about_ca_topic_score_codex":0.0031593374,"about_ca_topic_score_gemma":0.004679603,"teacher_disagreement_score":0.03197317,"about_ca_system_score_codex":0.0014894082,"about_ca_system_score_gemma":0.0030010296,"threshold_uncertainty_score":0.021660805},"labels":[],"label_agreement":null},{"id":"W2805097135","doi":"10.1109/icst.2018.00038","title":"Investigating NLP-Based Approaches for Predicting Manual Test Case Failure","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Test script; Regression testing; Test suite; Test (biology); Test case; Feature selection; Manual testing; Artificial intelligence; Heuristics; Test Management Approach; Feature (linguistics); Software regression; Machine learning; Software; Data mining; Programming language; Software system; Software quality; Software development; Regression analysis; Software construction","score_opus":0.057859450748061474,"score_gpt":0.2848859189331134,"score_spread":0.2270264681850519,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805097135","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20011373,0.0018757858,0.778304,0.0014200418,0.00009215346,0.00065853616,0.0027632427,0.010531013,0.0042415382],"genre_scores_gemma":[0.71185786,0.0005206754,0.28107637,0.00033611976,0.00010961211,0.0005094159,0.0040667974,0.0001998403,0.0013232398],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956449,0.0016494726,0.00041434547,0.0008995604,0.0011474439,0.0002442597],"domain_scores_gemma":[0.94451255,0.046356075,0.004087144,0.0012982062,0.0032760669,0.00046992986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004878044,0.0023502512,0.0012147236,0.00903872,0.00054856396,0.0016159966,0.0022300559,0.0015560685,0.0014793819],"category_scores_gemma":[0.027566373,0.00044021502,0.0010676943,0.0034005595,0.000647692,0.0024845588,0.00082303706,0.0017335773,0.00080839393],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046094402,0.0013366528,0.10039094,0.0011422605,0.00051069434,0.0007408017,0.00070277054,0.36741126,0.009053195,0.003244382,0.0049126935,0.51009333],"study_design_scores_gemma":[0.000019842208,0.00013465782,0.0066480096,0.000036180332,0.000047011566,0.00012013058,0.00015042638,0.9882578,0.0019798046,0.0019763543,0.0006002389,0.000029450299],"about_ca_topic_score_codex":0.017695531,"about_ca_topic_score_gemma":0.015920917,"teacher_disagreement_score":0.017695531,"about_ca_system_score_codex":0.0015778465,"about_ca_system_score_gemma":0.0016903712,"threshold_uncertainty_score":0.0351851},"labels":[],"label_agreement":null},{"id":"W2805382256","doi":"10.1145/3196831","title":"An Empirical Study of Meta- and Hyper-Heuristic Search for Multi-Objective Release Planning","year":2018,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Engineering and Physical Sciences Research Council; Dalian University of Technology; Core Research for Evolutional Science and Technology; Universiteit Utrecht; University College London","keywords":"Computer science; Heuristics; Heuristic; Meta heuristic; Variety (cybernetics); Genetic algorithm; Machine learning; Empirical research; Search algorithm; Beam search; Artificial intelligence; Mathematical optimization; Data mining; Algorithm; Mathematics; Statistics","score_opus":0.2711195274073803,"score_gpt":0.4358261630705214,"score_spread":0.1647066356631411,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805382256","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95105547,0.0046621608,0.035870552,0.00064538914,0.000050814986,0.0002758244,0.00074297236,0.00023262925,0.006464182],"genre_scores_gemma":[0.9621511,0.0006909062,0.035640642,0.00007679192,0.000018985413,0.00018311916,0.0007844206,0.000046278186,0.00040778145],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99240375,0.0056334836,0.00036740914,0.0005031058,0.000912673,0.00017956161],"domain_scores_gemma":[0.86168456,0.12571962,0.0038731783,0.004842163,0.0032499,0.0006306001],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012543626,0.00070341886,0.00057407084,0.001985681,0.0005259116,0.0012080122,0.0015280828,0.001098335,0.0013099618],"category_scores_gemma":[0.062257312,0.00040056085,0.0006340345,0.0034265595,0.00095560215,0.002743319,0.0006684618,0.0014636012,0.00018438257],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014895735,0.002280018,0.06607106,0.0014259561,0.0007526112,0.00018761338,0.0005051528,0.7375212,0.0014580645,0.008577608,0.0036579925,0.1760732],"study_design_scores_gemma":[0.00030313787,0.001971972,0.03452019,0.00022132172,0.00016537451,0.00028815048,0.00061661453,0.9518177,0.0020074588,0.004024442,0.0040091956,0.000054548298],"about_ca_topic_score_codex":0.0031093748,"about_ca_topic_score_gemma":0.004559703,"teacher_disagreement_score":0.012543626,"about_ca_system_score_codex":0.0013842066,"about_ca_system_score_gemma":0.0010689899,"threshold_uncertainty_score":0.06633788},"labels":[],"label_agreement":null},{"id":"W2806092851","doi":"10.1109/ase.2017.8115655","title":"Improved query reformulation for concept location using CodeRank and document structures","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Query expansion; Information retrieval; Source code; Query optimization; Sargable; Web query classification; Baseline (sea); Web search query; Task (project management); Software; Term (time); Query language; Code (set theory); Data mining; Search engine; Programming language","score_opus":0.028961210369557015,"score_gpt":0.3266036725533051,"score_spread":0.29764246218374807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806092851","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07473216,0.0021862432,0.89672744,0.0009839126,0.00019940922,0.0009046727,0.0020850678,0.019871324,0.0023097608],"genre_scores_gemma":[0.23751752,0.00069065363,0.74858123,0.00030522756,0.00025324294,0.00043618339,0.007777942,0.0007581014,0.0036798657],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946569,0.0017361388,0.0005438636,0.00079062436,0.0019797764,0.00029261783],"domain_scores_gemma":[0.9884731,0.0057688523,0.0008761123,0.0018265601,0.002826413,0.00022899022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029148227,0.0017046178,0.0021068216,0.006081227,0.0010292875,0.002045351,0.0018868389,0.0013136388,0.0040264307],"category_scores_gemma":[0.017869487,0.0005006887,0.001348797,0.004807763,0.000913522,0.005378274,0.0019580948,0.00204027,0.0026617923],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081626314,0.0007432863,0.003118756,0.0010260872,0.00014877794,0.00039092937,0.0012799244,0.037800238,0.07779412,0.013796489,0.03489578,0.82818943],"study_design_scores_gemma":[0.00034349822,0.0007705153,0.0027209185,0.0000656331,0.00020112842,0.0010472264,0.0010616706,0.87157303,0.07715976,0.014314325,0.030563496,0.00017877869],"about_ca_topic_score_codex":0.010732292,"about_ca_topic_score_gemma":0.010144599,"teacher_disagreement_score":0.010732292,"about_ca_system_score_codex":0.001475638,"about_ca_system_score_gemma":0.0030796358,"threshold_uncertainty_score":0.021339655},"labels":[],"label_agreement":null},{"id":"W2806282111","doi":"10.1002/stvr.1669","title":"MuMonDE: A framework for evaluating model clone detectors using model mutation analysis","year":2018,"lang":"en","type":"article","venue":"Software Testing Verification and Reliability","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Preprocessor; Mutation; Computer science; Mutation testing; Data mining; Detector; Precision and recall; Software engineering; Artificial intelligence; Genetics; Biology","score_opus":0.10310354306044087,"score_gpt":0.3755641998763948,"score_spread":0.2724606568159539,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806282111","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012272797,0.00009993945,0.9698473,0.00014300506,0.000021504104,0.00039591824,0.00042035713,0.015280713,0.0015184588],"genre_scores_gemma":[0.11326995,0.00006171018,0.8836408,0.000075488424,0.000014664864,0.00054187345,0.00094979943,0.00090699404,0.00053873623],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98712677,0.004957865,0.0014800511,0.0012855572,0.004701756,0.00044800248],"domain_scores_gemma":[0.96620876,0.020126276,0.0033394685,0.0043413965,0.0053999186,0.0005841756],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.017811961,0.0022560563,0.0014334898,0.0093310755,0.00089610304,0.004315079,0.0038385778,0.0014435882,0.0040845787],"category_scores_gemma":[0.056034375,0.0011020115,0.0025389884,0.0020239325,0.0020449795,0.003669194,0.0038627102,0.0019126841,0.0007796316],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080359966,0.0009598485,0.02600587,0.0011385577,0.0006939824,0.00050268177,0.001260492,0.39427823,0.025201248,0.124106735,0.0130271325,0.41202164],"study_design_scores_gemma":[0.000070077724,0.00020322403,0.0014997103,0.00013754432,0.00006534411,0.00013114873,0.00011178276,0.9632713,0.011154599,0.017317314,0.0059581664,0.00007977167],"about_ca_topic_score_codex":0.008010807,"about_ca_topic_score_gemma":0.008931389,"teacher_disagreement_score":0.98218805,"about_ca_system_score_codex":0.0027104998,"about_ca_system_score_gemma":0.002975256,"threshold_uncertainty_score":0.094199836},"labels":[],"label_agreement":null},{"id":"W2807405309","doi":"10.1145/3196883","title":"Inferring Extended Probabilistic Finite-State Automaton Models from Software Executions","year":2018,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Probabilistic automaton; Probabilistic logic; Automaton; Finite-state machine; Reinforcement learning; Decidability; Theoretical computer science; Software; Inference; Deterministic automaton; Flexibility (engineering); Büchi automaton; Deterministic finite automaton; Artificial intelligence; Programming language; Machine learning; Mathematics","score_opus":0.08968027278245541,"score_gpt":0.31906479794928544,"score_spread":0.22938452516683003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807405309","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031240245,0.00009004868,0.9652371,0.00013374393,0.000018364228,0.00009664472,0.0004263254,0.002219214,0.0005384465],"genre_scores_gemma":[0.5752295,0.0002798764,0.4205123,0.00010051342,0.000028984521,0.000353472,0.0021513149,0.00031771685,0.0010262796],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99762124,0.0008826413,0.00017098116,0.00053101976,0.00066611695,0.00012798498],"domain_scores_gemma":[0.9843507,0.012313972,0.00088488485,0.0016153094,0.00069169747,0.00014338916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002069041,0.0010192159,0.00060565706,0.0012984923,0.00038871018,0.0009670246,0.0020755434,0.0011448261,0.0015099411],"category_scores_gemma":[0.018179648,0.00076337124,0.001767405,0.0008255128,0.0010290806,0.0023919358,0.0013765214,0.0021466508,0.00043383843],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016723588,0.00017236985,0.008232334,0.00027939273,0.00013281302,0.00044991725,0.0004095875,0.877772,0.006483487,0.026711084,0.0009765938,0.07821313],"study_design_scores_gemma":[0.000007146339,0.000015471169,0.0002444067,0.0000138749865,0.000014241037,0.000028955126,0.000017170682,0.97711277,0.0019387917,0.02022257,0.00037622327,0.000008353268],"about_ca_topic_score_codex":0.005044392,"about_ca_topic_score_gemma":0.0112063745,"teacher_disagreement_score":0.005044392,"about_ca_system_score_codex":0.000972234,"about_ca_system_score_gemma":0.002041923,"threshold_uncertainty_score":0.01094228},"labels":[],"label_agreement":null},{"id":"W2807786268","doi":"10.24963/ijcai.2018/399","title":"Cutting the Software Building Efforts in Continuous Integration by Semi-Supervised Online AUC Optimization","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Novelis (Canada)","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"Computer science; Outcome (game theory); Software; Event (particle physics); Code (set theory); Resource (disambiguation); Software engineering; Data science; Machine learning; Artificial intelligence; Set (abstract data type)","score_opus":0.012350041946344934,"score_gpt":0.26505419219190113,"score_spread":0.2527041502455562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807786268","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.089467965,0.001523291,0.8976964,0.000874787,0.00012073868,0.0002217463,0.00042794438,0.0075358795,0.0021312956],"genre_scores_gemma":[0.78711575,0.0003683057,0.2031048,0.000821767,0.00028972593,0.0005013961,0.002970239,0.0007334062,0.0040946435],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99693155,0.0009973084,0.00022041073,0.000988463,0.00055076234,0.00031148447],"domain_scores_gemma":[0.98551536,0.008777571,0.0012943689,0.0013222351,0.0023545113,0.00073601573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042800675,0.0024981746,0.0027746903,0.0017708349,0.0009133727,0.0016401456,0.004064204,0.0026480074,0.001749368],"category_scores_gemma":[0.014865429,0.0011039712,0.0012725664,0.0015821583,0.0016453656,0.0030577278,0.0021328644,0.0039324868,0.0011097345],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000577574,0.00072233944,0.008013638,0.00034854613,0.00018739804,0.0002202578,0.00022321487,0.6832829,0.0037607697,0.002251607,0.011103819,0.28930795],"study_design_scores_gemma":[0.000011151533,0.000047742757,0.00034210642,0.00000788497,0.000009362928,0.000021334072,0.000009994845,0.99752265,0.00058393594,0.0011822927,0.0002541643,0.000007442343],"about_ca_topic_score_codex":0.0076418514,"about_ca_topic_score_gemma":0.008041951,"teacher_disagreement_score":0.0076418514,"about_ca_system_score_codex":0.0012784598,"about_ca_system_score_gemma":0.0022886328,"threshold_uncertainty_score":0.0226354},"labels":[],"label_agreement":null},{"id":"W2808113972","doi":"10.1145/3183519.3183547","title":"An experience report on defect modelling in practice","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":125,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Field (mathematics); Key (lock); Work (physics); Data science; Computer science; Engineering ethics; Management science; Knowledge management; Process management; Engineering; Computer security","score_opus":0.04393555895126731,"score_gpt":0.3576458420899845,"score_spread":0.3137102831387172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808113972","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4261788,0.017810943,0.34738252,0.03209763,0.0009227368,0.0011856575,0.001816345,0.0068528447,0.16575246],"genre_scores_gemma":[0.7713235,0.010560582,0.17822196,0.0021313042,0.00018300899,0.00048446126,0.0029077157,0.0024739553,0.03171361],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9479838,0.020831708,0.0042827986,0.0050174426,0.020345552,0.0015387256],"domain_scores_gemma":[0.8482264,0.08080692,0.010047397,0.026938342,0.029661512,0.004319432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.042986613,0.0010867304,0.00072880945,0.0066895396,0.0025284197,0.008980042,0.0052723633,0.002709508,0.008413284],"category_scores_gemma":[0.1505509,0.0008932614,0.0012103671,0.0070888326,0.0036662593,0.010542741,0.0079122195,0.0028651608,0.0034711687],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021758345,0.00093561446,0.081369266,0.0013751881,0.000112558955,0.0021463803,0.08557606,0.009890298,0.0027211877,0.040579233,0.046375815,0.7287008],"study_design_scores_gemma":[0.00007166323,0.001179866,0.036601286,0.004606051,0.00020397827,0.009674844,0.065561466,0.04149492,0.0065713543,0.038374543,0.79525685,0.0004032266],"about_ca_topic_score_codex":0.01025405,"about_ca_topic_score_gemma":0.009756874,"teacher_disagreement_score":0.042986613,"about_ca_system_score_codex":0.0059005595,"about_ca_system_score_gemma":0.0065589743,"threshold_uncertainty_score":0.22733766},"labels":[],"label_agreement":null},{"id":"W2808372290","doi":"10.1145/3180155.3182535","title":"Analyzing the effects of test driven development in GitHub","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Agile software development; Test-driven development; Computer science; Commit; Software engineering; Software quality; Software; Quality (philosophy); Software development; Software development process; Operating system; Database","score_opus":0.009733605887843307,"score_gpt":0.25066111033065513,"score_spread":0.24092750444281183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808372290","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9973024,0.00012756919,0.0012088568,0.00007743234,0.000011155016,0.00015269037,0.00033585072,0.00008978302,0.0006941418],"genre_scores_gemma":[0.99472946,0.00006777179,0.0033367258,0.000104010986,0.00002253363,0.00049964205,0.0008093735,0.000059802005,0.00037075422],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98449534,0.0072611924,0.0008486503,0.0020900632,0.00444666,0.00085807586],"domain_scores_gemma":[0.8131115,0.1305515,0.02485484,0.017998628,0.010511701,0.0029718673],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010918248,0.00066900614,0.0005183462,0.0028090538,0.0005355074,0.0015898546,0.0018560454,0.0007626423,0.001193257],"category_scores_gemma":[0.08308056,0.0004601357,0.0011560731,0.0031073862,0.0017963372,0.0015858246,0.0018830881,0.001676588,0.00036783877],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.014033106,0.008576002,0.7175907,0.0018959879,0.0018786696,0.00087021576,0.0070974734,0.042503804,0.025679423,0.0029860626,0.003984326,0.17290416],"study_design_scores_gemma":[0.0002617979,0.007289551,0.95268863,0.00010391973,0.0003769609,0.00019150438,0.0014993254,0.02661694,0.008275909,0.00074131513,0.0018487512,0.00010534464],"about_ca_topic_score_codex":0.011943473,"about_ca_topic_score_gemma":0.010025911,"teacher_disagreement_score":0.98908174,"about_ca_system_score_codex":0.004215174,"about_ca_system_score_gemma":0.0015458638,"threshold_uncertainty_score":0.05774188},"labels":[],"label_agreement":null},{"id":"W2808509424","doi":"10.24963/ijcai.2018/394","title":"Positive and Unlabeled Learning for Detecting Software Functional Clones with Adversarial Training","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Novelis (Canada)","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"clone (Java method); Computer science; Adversarial system; Software maintenance; Robustness (evolution); Software; Artificial intelligence; Machine learning; Task (project management); Software development; Software engineering; Programming language; Engineering; Biology; Gene; Genetics","score_opus":0.024836931718548633,"score_gpt":0.2511193994927301,"score_spread":0.22628246777418146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808509424","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07983759,0.0006971031,0.9148811,0.0005789291,0.00008290886,0.00011421962,0.0001749392,0.0016215771,0.0020116272],"genre_scores_gemma":[0.8568934,0.0002116043,0.13832492,0.000671585,0.000114807444,0.00022934825,0.00088281464,0.00013796959,0.0025335539],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983525,0.000674,0.00006808802,0.00044741872,0.0003377355,0.00012019446],"domain_scores_gemma":[0.98827153,0.009216232,0.00057272054,0.00096305204,0.00074178196,0.00023472797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003283854,0.0014661435,0.0011671147,0.0010528433,0.00061725243,0.0006623381,0.0026807815,0.002099628,0.0012912313],"category_scores_gemma":[0.011715138,0.0004952831,0.00075708405,0.0006168549,0.0022270384,0.0018875324,0.001931539,0.002396186,0.00037534517],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036607785,0.00021472083,0.0049913214,0.00015373676,0.00009071793,0.000218879,0.00016956036,0.84193265,0.004660834,0.009520009,0.0038509003,0.1338305],"study_design_scores_gemma":[0.000004947617,0.000022376254,0.000121246936,0.0000047515928,0.000005743618,0.000018196295,0.000005809156,0.99572515,0.0007353157,0.003196377,0.00015622909,0.000003900083],"about_ca_topic_score_codex":0.0030933102,"about_ca_topic_score_gemma":0.002981377,"teacher_disagreement_score":0.003283854,"about_ca_system_score_codex":0.00102044,"about_ca_system_score_gemma":0.0009123303,"threshold_uncertainty_score":0.017366886},"labels":[],"label_agreement":null},{"id":"W2808519529","doi":"10.1145/3180155.3182522","title":"On the use of hidden Markov model to predict the time to fix bugs","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Hidden Markov model; Software bug; Computer science; Software regression; Software; Software engineering; Data science; Software development; Data mining; Artificial intelligence; Software quality; Programming language","score_opus":0.03886277116275344,"score_gpt":0.26186448972258713,"score_spread":0.2230017185598337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808519529","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26029527,0.0017167424,0.730402,0.0013112886,0.00018231814,0.00013239296,0.0008254537,0.0022190274,0.0029154492],"genre_scores_gemma":[0.8962263,0.0010838418,0.09866925,0.00021784786,0.00013141749,0.00011748738,0.0013408405,0.000096762655,0.0021162478],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99916565,0.0003574188,0.00006549988,0.00021032566,0.00010876414,0.00009233758],"domain_scores_gemma":[0.98620343,0.012323962,0.000509969,0.0002351975,0.00055752724,0.00016997005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035747357,0.0012305891,0.0009482644,0.0021923836,0.0006036432,0.0011804535,0.0009517618,0.0012424475,0.0014914222],"category_scores_gemma":[0.011670043,0.00073540566,0.0012466224,0.0014290954,0.00038328997,0.001798649,0.0006262011,0.0015102352,0.0006215231],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005017253,0.0004908795,0.05013271,0.00012989057,0.000353286,0.0002295664,0.000286422,0.77542436,0.0015943868,0.0039904495,0.0021826318,0.1646837],"study_design_scores_gemma":[0.000006598929,0.000035261935,0.0013777873,0.000011296873,0.000018351295,0.000015827214,0.000013197248,0.9969739,0.00014813864,0.0012920563,0.0000964881,0.0000112046],"about_ca_topic_score_codex":0.046193928,"about_ca_topic_score_gemma":0.04204312,"teacher_disagreement_score":0.046193928,"about_ca_system_score_codex":0.0009491104,"about_ca_system_score_gemma":0.0016052941,"threshold_uncertainty_score":0.0918501},"labels":[],"label_agreement":null},{"id":"W2808569344","doi":"10.1145/3183399.3183418","title":"Which library should I use?","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Task (project management); Reuse; Software; World Wide Web; Metric (unit); Quality (philosophy); Code (set theory); Software engineering; Usability; Data science; Human–computer interaction; Engineering; Set (abstract data type)","score_opus":0.04754281956395838,"score_gpt":0.2831562047688548,"score_spread":0.2356133852048964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808569344","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30330813,0.007802608,0.09700967,0.08599015,0.003184165,0.0013105222,0.017587598,0.028547,0.4552602],"genre_scores_gemma":[0.73401415,0.006708219,0.09029615,0.01000682,0.0012771611,0.0007938374,0.011688129,0.008875462,0.1363401],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9947931,0.0016517866,0.0004913352,0.00065099035,0.0018950374,0.0005177637],"domain_scores_gemma":[0.9712106,0.0070262696,0.003595635,0.002833649,0.009969727,0.0053642727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044538206,0.0006057234,0.00076768803,0.0043373997,0.0015954564,0.0075917384,0.0013512559,0.0015560167,0.02138396],"category_scores_gemma":[0.05064483,0.0004993561,0.00059459865,0.0048789075,0.0011965636,0.012700907,0.0020695254,0.0016106856,0.023849355],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005164,0.00044580954,0.12718563,0.0011375842,0.000150473,0.00082810834,0.007699891,0.00039972784,0.0035059517,0.017782316,0.28743875,0.5529093],"study_design_scores_gemma":[0.00007102614,0.00040066638,0.08723873,0.0009964651,0.00020418936,0.0030221366,0.012268964,0.0029992945,0.008466703,0.010612969,0.87335676,0.0003620998],"about_ca_topic_score_codex":0.006660878,"about_ca_topic_score_gemma":0.008501017,"teacher_disagreement_score":0.02138396,"about_ca_system_score_codex":0.0021681418,"about_ca_system_score_gemma":0.0023450346,"threshold_uncertainty_score":0.07153648},"labels":[],"label_agreement":null},{"id":"W2808958974","doi":"10.1145/3183440.3194972","title":"ALPACA-advanced linguistic pattern and concept analysis framework for software engineering corpora","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software engineering; Software; Natural language processing; Artificial intelligence; Programming language","score_opus":0.015372292295679678,"score_gpt":0.2825382621082936,"score_spread":0.2671659698126139,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808958974","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021165446,0.0002027105,0.91622543,0.0003914512,0.00009746082,0.0010033922,0.014428915,0.061632685,0.0039014192],"genre_scores_gemma":[0.01570875,0.00016774505,0.9523743,0.00017026742,0.00009178653,0.002069905,0.022746164,0.003924855,0.0027462244],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99429494,0.0016921195,0.0008987233,0.0014130721,0.0014689098,0.00023222785],"domain_scores_gemma":[0.98761034,0.005589278,0.001234846,0.0019939754,0.003085961,0.00048551738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007842713,0.0018507986,0.0013274343,0.011390503,0.0022175037,0.005947767,0.0025478448,0.0014090033,0.024900645],"category_scores_gemma":[0.02168953,0.0017234152,0.0030672697,0.0077582197,0.0015762106,0.006785322,0.0045457473,0.0040954286,0.013453971],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007317936,0.0003135082,0.004352662,0.0032561605,0.00047686318,0.0017331634,0.0053168316,0.013646375,0.026334101,0.16175006,0.1805887,0.60149974],"study_design_scores_gemma":[0.00027842206,0.00015485633,0.0049634287,0.00043403785,0.00013310855,0.0016303816,0.0012289655,0.18885097,0.013579487,0.12677868,0.6617136,0.00025396314],"about_ca_topic_score_codex":0.008471006,"about_ca_topic_score_gemma":0.011382262,"teacher_disagreement_score":0.024900645,"about_ca_system_score_codex":0.0019463273,"about_ca_system_score_gemma":0.004578669,"threshold_uncertainty_score":0.08330095},"labels":[],"label_agreement":null},{"id":"W2809042601","doi":"10.1145/3183440.3195003","title":"Improving bug localization with report quality dynamics and query reformulation","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Quality (philosophy); Software bug; Dynamics (music); Empirical research; Information retrieval; Stack (abstract data type); Data mining; Artificial intelligence; Software; Programming language","score_opus":0.02016616833125125,"score_gpt":0.29885155190049356,"score_spread":0.27868538356924233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2809042601","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7578231,0.0053333514,0.20048775,0.0019793801,0.0001552849,0.0009673113,0.0011740644,0.029093994,0.002985704],"genre_scores_gemma":[0.86915356,0.00052520935,0.1264255,0.00029595735,0.00015150529,0.0002762675,0.0015685285,0.00092417235,0.0006793435],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9698183,0.01342866,0.0034674972,0.0046685846,0.0074975053,0.0011194402],"domain_scores_gemma":[0.75694656,0.16781661,0.027714856,0.028442105,0.017193202,0.0018866342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029630456,0.0020528906,0.0021460545,0.005827591,0.00080612395,0.0038786193,0.0032551836,0.0015159934,0.0016655148],"category_scores_gemma":[0.22483644,0.0008966244,0.0010851765,0.004587137,0.0012531094,0.0069654216,0.0025390463,0.0025199926,0.00097171724],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023959999,0.0025385288,0.12044125,0.0033806143,0.00063803664,0.00056026125,0.0075887614,0.044359747,0.08443724,0.003003986,0.009248306,0.7214073],"study_design_scores_gemma":[0.0011088959,0.006726852,0.14800355,0.00051868777,0.0022927453,0.0025365131,0.0045892093,0.7058767,0.10238426,0.007711852,0.017607037,0.00064362184],"about_ca_topic_score_codex":0.0046668877,"about_ca_topic_score_gemma":0.0024567249,"teacher_disagreement_score":0.029630456,"about_ca_system_score_codex":0.0014911303,"about_ca_system_score_gemma":0.0025395767,"threshold_uncertainty_score":0.1567027},"labels":[],"label_agreement":null},{"id":"W2809153102","doi":"10.1145/3197091.3205832","title":"Code reviews in large, first-year courses","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Code (set theory); Implementation; Code review; Programming style; Software engineering; Programming language; Mathematics education; Multimedia; Software; Static program analysis; Software development; Psychology","score_opus":0.03395933909286687,"score_gpt":0.31982509606048615,"score_spread":0.2858657569676193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2809153102","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7748907,0.0036275901,0.10757325,0.01865197,0.00494707,0.0058327317,0.0012082007,0.016386675,0.06688189],"genre_scores_gemma":[0.7354517,0.0016873502,0.18625408,0.005602728,0.0015808443,0.0031824543,0.002289232,0.0038050383,0.060146533],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9669001,0.012845444,0.0019176423,0.004712483,0.011364555,0.0022597376],"domain_scores_gemma":[0.687646,0.1182698,0.032043476,0.02477844,0.093522154,0.0437401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028164338,0.001235898,0.0011974112,0.0041458774,0.00457149,0.0052514207,0.0039102077,0.0018767723,0.010889768],"category_scores_gemma":[0.16769744,0.0010695262,0.00063382037,0.0026189561,0.0014737219,0.0035386262,0.005005962,0.0035727592,0.008136475],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006788537,0.004080099,0.028029984,0.0008146692,0.000077832774,0.00094186165,0.01416307,0.0019607702,0.016065992,0.002372021,0.16716397,0.7636509],"study_design_scores_gemma":[0.00087543903,0.0076812645,0.16779217,0.0021027387,0.00015431506,0.0021907499,0.024412738,0.014794868,0.040733486,0.018427245,0.72015685,0.00067810173],"about_ca_topic_score_codex":0.0014666624,"about_ca_topic_score_gemma":0.0075242994,"teacher_disagreement_score":0.028164338,"about_ca_system_score_codex":0.0054542674,"about_ca_system_score_gemma":0.009410642,"threshold_uncertainty_score":0.14894909},"labels":[],"label_agreement":null},{"id":"W2809245278","doi":"10.1145/3183440.3195005","title":"Fast, scalable and user-guided clone detection","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Scalability; clone (Java method); Operating system; Biology; Gene","score_opus":0.018363515250642736,"score_gpt":0.2662486519695933,"score_spread":0.24788513671895057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2809245278","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0416298,0.0005392956,0.82151526,0.0003416208,0.00008422328,0.00031176096,0.0021497335,0.13170376,0.0017246074],"genre_scores_gemma":[0.18832989,0.0002742562,0.7943197,0.00023171051,0.000069567264,0.0003207265,0.0081985565,0.0046492545,0.0036062829],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9948218,0.00064030185,0.00040536246,0.0012891977,0.0024993315,0.0003439976],"domain_scores_gemma":[0.9809241,0.0060572033,0.001604875,0.005292225,0.005516603,0.00060494523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035951277,0.0014472504,0.001389802,0.0046788724,0.00084527646,0.0031641044,0.0026101922,0.0015001578,0.0018979928],"category_scores_gemma":[0.024652703,0.0007769061,0.0011685637,0.0037153983,0.0006880246,0.004257485,0.0032455856,0.0014836865,0.0028117476],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007840225,0.0003296034,0.029003374,0.0005786795,0.00024341658,0.00059629534,0.0010371922,0.0075717927,0.1285916,0.004776044,0.04232972,0.7841584],"study_design_scores_gemma":[0.00020933065,0.0004080018,0.026560321,0.00011211675,0.00018232115,0.0021959676,0.0006873226,0.6041945,0.28577903,0.024441538,0.054959673,0.000269762],"about_ca_topic_score_codex":0.003435313,"about_ca_topic_score_gemma":0.00493498,"teacher_disagreement_score":0.0046788724,"about_ca_system_score_codex":0.0007319113,"about_ca_system_score_gemma":0.0013998976,"threshold_uncertainty_score":0.019013107},"labels":[],"label_agreement":null},{"id":"W2810627707","doi":"10.1007/s10664-018-9634-5","title":"How do developers utilize source code from stack overflow?","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":129,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Source code; Copying; Codebase; Code reuse; Reuse; Code review; Programming language; Code (set theory); Stack (abstract data type); Software engineering; World Wide Web; Operating system; Static program analysis; Software; Software development; Engineering; Set (abstract data type)","score_opus":0.03258012208955109,"score_gpt":0.27706487837659804,"score_spread":0.24448475628704697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2810627707","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9576204,0.0008502946,0.022813613,0.0040575853,0.00009546603,0.0002550305,0.0005785972,0.0022729319,0.011456121],"genre_scores_gemma":[0.96653,0.001053503,0.019524798,0.0011638462,0.000058496113,0.0002727035,0.0015550312,0.0022516674,0.007589886],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9830078,0.005370624,0.0012919918,0.0019796395,0.0071321046,0.0012178249],"domain_scores_gemma":[0.8591958,0.08143572,0.020903405,0.013463951,0.022463614,0.0025374452],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.016612967,0.00074211456,0.00044804285,0.00627989,0.0023224752,0.004476807,0.0014784412,0.0017224196,0.00217407],"category_scores_gemma":[0.17544097,0.0010115751,0.00061381405,0.0040213913,0.0028559975,0.012467347,0.0042458773,0.001496183,0.0014599632],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034679435,0.00034267074,0.28949273,0.0010266607,0.00011724637,0.003472979,0.36116168,0.0009053434,0.010014689,0.005389312,0.015231742,0.31249812],"study_design_scores_gemma":[0.00011519407,0.0008215495,0.47661754,0.0030917965,0.0002780044,0.0069891536,0.21481434,0.00983843,0.021569867,0.012170052,0.25308362,0.0006104334],"about_ca_topic_score_codex":0.010385393,"about_ca_topic_score_gemma":0.011947373,"teacher_disagreement_score":0.98338705,"about_ca_system_score_codex":0.002656428,"about_ca_system_score_gemma":0.004246478,"threshold_uncertainty_score":0.087858796},"labels":[],"label_agreement":null},{"id":"W2810784885","doi":"10.1145/3167132.3167292","title":"A model for analysis and presentation of design pattern detection results","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Presentation (obstetrics); Comprehension; Process (computing); Software design pattern; Program comprehension; Software engineering; Software; Software design; Software system; Artificial intelligence; Human–computer interaction; Software development; Programming language","score_opus":0.06932823813603578,"score_gpt":0.32102712610589046,"score_spread":0.2516988879698547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2810784885","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011444729,0.00022172052,0.99005556,0.001049581,0.000080008474,0.00023550449,0.00044226128,0.002911135,0.003859777],"genre_scores_gemma":[0.07132604,0.00090699265,0.91472167,0.0007922036,0.00031399936,0.0016369208,0.0017762186,0.0008285936,0.007697408],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9865416,0.004665406,0.0016591115,0.0022404394,0.004237736,0.0006556127],"domain_scores_gemma":[0.9505916,0.031808008,0.003194061,0.007607416,0.006134437,0.0006644785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014317784,0.0032465511,0.0016696292,0.0075060185,0.0014297154,0.0120965615,0.006117815,0.00595605,0.017099777],"category_scores_gemma":[0.060100626,0.001676276,0.0044032014,0.005397848,0.003687886,0.016435985,0.0037956652,0.0043388973,0.009578502],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005686966,0.00030660862,0.0028167386,0.00097594067,0.00024511197,0.0013055016,0.0018975365,0.07136317,0.0061341906,0.75767887,0.018528678,0.13817888],"study_design_scores_gemma":[0.000103791055,0.00017645564,0.0003584278,0.00026754168,0.00013820823,0.00056644896,0.00018213908,0.4420988,0.0031812065,0.51632786,0.03649927,0.0000998382],"about_ca_topic_score_codex":0.0044907397,"about_ca_topic_score_gemma":0.0024844846,"teacher_disagreement_score":0.017099777,"about_ca_system_score_codex":0.003349742,"about_ca_system_score_gemma":0.0035903347,"threshold_uncertainty_score":0.07572055},"labels":[],"label_agreement":null},{"id":"W2812765312","doi":"10.1111/cgf.13425","title":"ThreadReconstructor: Modeling Reply‐Chains to Untangle Conversational Text through Visual Analytics","year":2018,"lang":"en","type":"article","venue":"Computer Graphics Forum","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Heuristics; Visual analytics; Visualization; Analytics; Human–computer interaction; Data visualization; Artificial intelligence; Data science; Machine learning; Information retrieval","score_opus":0.03116623171626524,"score_gpt":0.2922717656578509,"score_spread":0.26110553394158564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2812765312","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04204684,0.00021890654,0.9376605,0.0003036409,0.0000471016,0.00020445978,0.0014560996,0.016638134,0.0014242802],"genre_scores_gemma":[0.3089071,0.00016734535,0.6844537,0.000081641185,0.00003949439,0.00042329918,0.0025471216,0.0018176378,0.0015627203],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99856335,0.00064729573,0.00008484334,0.0003412395,0.00026447067,0.00009872251],"domain_scores_gemma":[0.991867,0.0057982923,0.0006940124,0.00074838166,0.0005956027,0.00029669734],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034922613,0.0013302952,0.0006471873,0.003896906,0.0006224242,0.0032578877,0.001681618,0.0010625185,0.0058571696],"category_scores_gemma":[0.014577679,0.0006004955,0.0010382249,0.001525628,0.0008808227,0.00295277,0.0022309292,0.0014175187,0.001618844],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019420924,0.00050445844,0.025196105,0.0015364432,0.00026329575,0.0006650246,0.018075591,0.16760388,0.055846903,0.05259868,0.021442244,0.65432525],"study_design_scores_gemma":[0.000030062703,0.00006190035,0.0014976413,0.00007021001,0.00001736226,0.0000583929,0.00063205045,0.95952,0.008098783,0.021428186,0.00854338,0.000042054937],"about_ca_topic_score_codex":0.005267274,"about_ca_topic_score_gemma":0.0053337608,"teacher_disagreement_score":0.0058571696,"about_ca_system_score_codex":0.0008636198,"about_ca_system_score_gemma":0.001136176,"threshold_uncertainty_score":0.019594193},"labels":[],"label_agreement":null},{"id":"W2883027691","doi":"10.1145/3196398.3196457","title":"Developer interaction traces backed by IDE screen recordings from think aloud sessions","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Concordia University","funders":"Polytechnique Montréal","keywords":"Correctness; Computer science; Program comprehension; Set (abstract data type); Human–computer interaction; Empirical research; Data science; Comprehension; Ground truth; Artificial intelligence; Software; Programming language; Software system","score_opus":0.024110489005982356,"score_gpt":0.29143017730406157,"score_spread":0.2673196882980792,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883027691","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8355318,0.0006110426,0.05735648,0.000603292,0.0002035362,0.0009529534,0.08578837,0.0033855157,0.015567089],"genre_scores_gemma":[0.86106527,0.0005338806,0.063078105,0.00012620135,0.00012149833,0.0016840397,0.06683143,0.0006065342,0.0059530074],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9975998,0.00061425613,0.0002542671,0.00037496476,0.0009794086,0.00017728559],"domain_scores_gemma":[0.9684542,0.017728671,0.0029757079,0.00410678,0.0057498887,0.0009847468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015563576,0.00070494454,0.00047431525,0.005374329,0.00055752497,0.0010096629,0.0005709768,0.0007786049,0.0033100527],"category_scores_gemma":[0.021461898,0.00035106204,0.00024626948,0.0044746567,0.0005388632,0.0011140895,0.001185744,0.0011944717,0.0023137906],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001963227,0.0018774782,0.25048226,0.0043020886,0.00030752228,0.00270279,0.0534493,0.0041143945,0.056015704,0.0042524594,0.07778579,0.542747],"study_design_scores_gemma":[0.00016808292,0.0012482246,0.7774875,0.0010965861,0.0001268794,0.0020066777,0.027841343,0.017302405,0.04390669,0.007454375,0.12110439,0.00025690123],"about_ca_topic_score_codex":0.00371751,"about_ca_topic_score_gemma":0.008474939,"teacher_disagreement_score":0.005374329,"about_ca_system_score_codex":0.00036320934,"about_ca_system_score_gemma":0.00059758936,"threshold_uncertainty_score":0.011073232},"labels":[],"label_agreement":null},{"id":"W2883195769","doi":"10.1109/icsme.2018.00057","title":"Effective Reformulation of Query for Code Search Using Crowdsourced Knowledge and Extra-Large Data Analytics","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Information retrieval; Code (set theory); Search engine; Web search query; Semantic search; Relevance (law); Query expansion; Programming language","score_opus":0.13152020755385394,"score_gpt":0.3966857943397065,"score_spread":0.26516558678585256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883195769","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.077710345,0.0015862146,0.89138484,0.002611485,0.000201393,0.0013332888,0.0046251793,0.014627264,0.005920085],"genre_scores_gemma":[0.46451366,0.0004735758,0.51635414,0.0008704596,0.00022744281,0.0010104646,0.012269624,0.00072330853,0.0035573547],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98675865,0.0051215184,0.00091879175,0.0027968662,0.0037867555,0.0006173767],"domain_scores_gemma":[0.97825634,0.01217088,0.0011662832,0.0044362647,0.0033586342,0.0006115693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00726889,0.0020142435,0.0026406718,0.008596272,0.0018911329,0.0032801805,0.003229929,0.0024530233,0.0033002915],"category_scores_gemma":[0.033913027,0.0006919773,0.0016218339,0.0068070963,0.0020393955,0.006490489,0.0060240985,0.0020624646,0.0023599605],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013637526,0.0014141368,0.009458874,0.0025842236,0.00034393132,0.0010231291,0.005461299,0.08383508,0.04784443,0.02652795,0.061024785,0.75911844],"study_design_scores_gemma":[0.0002319171,0.00035666302,0.0034875928,0.000112705966,0.00014170578,0.00034325733,0.00283853,0.89386714,0.017066585,0.058353968,0.023052836,0.00014712042],"about_ca_topic_score_codex":0.021572124,"about_ca_topic_score_gemma":0.025258414,"teacher_disagreement_score":0.021572124,"about_ca_system_score_codex":0.0025579322,"about_ca_system_score_gemma":0.0054730964,"threshold_uncertainty_score":0.04289311},"labels":[],"label_agreement":null},{"id":"W2883229686","doi":"10.1145/3196398.3196469","title":"Revisiting \"programmers' build errors\" in the visual studio context","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Context (archaeology); Studio; Programming language; Human–computer interaction; Computer graphics (images); Microsoft Visual Studio; Software engineering; Engineering drawing; Multimedia; Software; Engineering","score_opus":0.02762530743136254,"score_gpt":0.32899214899351564,"score_spread":0.3013668415621531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883229686","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97721964,0.0014924647,0.010476265,0.0019367767,0.000106528314,0.00006299915,0.0008925313,0.00028194115,0.0075308383],"genre_scores_gemma":[0.99110764,0.00040847505,0.0056532547,0.00050406443,0.00009392552,0.00004862282,0.00083043706,0.00027503865,0.0010786299],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97808814,0.009144446,0.0016545447,0.0036077213,0.0062044067,0.0013008132],"domain_scores_gemma":[0.7911563,0.13299267,0.042240188,0.015119549,0.015954338,0.0025370207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014858825,0.0006428289,0.0004370229,0.006096386,0.0015655243,0.0040547308,0.0011964884,0.0010260699,0.0016009856],"category_scores_gemma":[0.118490316,0.0004985012,0.0005139587,0.005946408,0.0023904457,0.0057608476,0.004012688,0.002384921,0.00042311227],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024716274,0.00018557502,0.8481268,0.0006030904,0.00013811243,0.00078466575,0.057904106,0.0011063139,0.0017332551,0.0036067867,0.007171103,0.078392945],"study_design_scores_gemma":[0.000023781506,0.00022403864,0.9122403,0.0005334697,0.00013321839,0.0011567048,0.046766777,0.0051599112,0.002044006,0.0062447917,0.025392046,0.0000809987],"about_ca_topic_score_codex":0.012783343,"about_ca_topic_score_gemma":0.030785477,"teacher_disagreement_score":0.014858825,"about_ca_system_score_codex":0.001271585,"about_ca_system_score_gemma":0.0018714517,"threshold_uncertainty_score":0.07858193},"labels":[],"label_agreement":null},{"id":"W2883393561","doi":"10.1002/smr.1965","title":"Program comprehension through reverse‐engineered sequence diagrams: A systematic review","year":2018,"lang":"en","type":"review","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Sequence diagram; Program comprehension; Documentation; Reverse engineering; Sequence (biology); Process (computing); Set (abstract data type); Context (archaeology); Software engineering; Software; Comprehension; Use Case Diagram; Data science; Information retrieval; Programming language; Data mining; Unified Modeling Language; Software system; Class diagram","score_opus":0.06289656528401565,"score_gpt":0.3715168890266922,"score_spread":0.30862032374267656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883393561","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006915924,0.9972313,0.0008430897,0.00029009,0.00009073253,0.0001928412,0.00019169228,0.000018810215,0.0004499168],"genre_scores_gemma":[0.0055540744,0.9915951,0.001900361,0.00025736087,0.000037208716,0.00030101647,0.00020178227,0.000009072762,0.00014391355],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9946826,0.0020588215,0.0015826286,0.00046300257,0.0011017909,0.00011114457],"domain_scores_gemma":[0.96310854,0.027768834,0.0035997208,0.00070093287,0.0045422637,0.00027971348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00903588,0.0012876882,0.0032367215,0.014251777,0.0005268237,0.0018501969,0.0024469274,0.0013802325,0.004020784],"category_scores_gemma":[0.04593365,0.00080068613,0.00281817,0.01106824,0.0008316427,0.0038679026,0.0014420974,0.0010995655,0.0005919399],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000098639066,0.000044834243,0.0005554994,0.66126466,0.0016210387,0.00015578067,0.0006430376,0.00032299277,0.0004519246,0.0015686561,0.005249529,0.3280235],"study_design_scores_gemma":[0.00010449023,0.00026624586,0.0027086257,0.80297273,0.011874428,0.0007383815,0.00082667917,0.00026493863,0.0007894903,0.0018577103,0.17752527,0.000070942275],"about_ca_topic_score_codex":0.0045949016,"about_ca_topic_score_gemma":0.012562635,"teacher_disagreement_score":0.014251777,"about_ca_system_score_codex":0.0018510347,"about_ca_system_score_gemma":0.012166226,"threshold_uncertainty_score":0.04778689},"labels":[],"label_agreement":null},{"id":"W2883473889","doi":"10.1145/3196398.3196471","title":"Do software engineers use autocompletion features differently than other developers?","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Software engineering; Software; Programming language","score_opus":0.030983112206759155,"score_gpt":0.26913356508258657,"score_spread":0.23815045287582742,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883473889","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9945146,0.00039433976,0.0018001848,0.00061298686,0.000023851915,0.000023294877,0.0007566224,0.00007254889,0.0018017022],"genre_scores_gemma":[0.9954527,0.00027466562,0.0012256851,0.0003427327,0.000023592307,0.00004643114,0.0013126103,0.00006616892,0.0012554015],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9946104,0.0016717481,0.000585457,0.0012767352,0.0013584271,0.0004972198],"domain_scores_gemma":[0.9220654,0.045290116,0.019940214,0.0049036434,0.0055485256,0.0022520584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052102725,0.00027612,0.0004422383,0.0023605565,0.0005654584,0.001344649,0.0005607327,0.00081264845,0.0023324513],"category_scores_gemma":[0.052721977,0.0003078625,0.00032447142,0.0021258015,0.0006841844,0.0024011082,0.0008058905,0.00070059625,0.0012581672],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014641171,0.00005960867,0.95611537,0.00010764357,0.000076554454,0.0001164524,0.007436638,0.00005689966,0.0007869238,0.00018231054,0.0016183018,0.033297],"study_design_scores_gemma":[0.00001798297,0.00010112532,0.9792762,0.00008153678,0.00005489459,0.00066784606,0.010430112,0.00060577015,0.0012472487,0.0007540068,0.0067303753,0.000032957636],"about_ca_topic_score_codex":0.0029597767,"about_ca_topic_score_gemma":0.0074623134,"teacher_disagreement_score":0.0052102725,"about_ca_system_score_codex":0.00035785892,"about_ca_system_score_gemma":0.00039670544,"threshold_uncertainty_score":0.02755487},"labels":[],"label_agreement":null},{"id":"W2883496053","doi":"10.1145/3194095.3194101","title":"Towards a classification of bugs to facilitate software maintainability tasks","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Software bug; Computer science; Maintainability; Software maintenance; Process (computing); Software quality; Software; Software regression; Set (abstract data type); Code (set theory); Software engineering; Software development; Programming language","score_opus":0.0572385695500097,"score_gpt":0.30362185949951653,"score_spread":0.24638328994950684,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883496053","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.121386856,0.0046555744,0.8240757,0.0050898436,0.00073079055,0.0026988462,0.012240791,0.01898708,0.010134509],"genre_scores_gemma":[0.112747975,0.0012142826,0.86304665,0.00049095694,0.0002849601,0.00095713144,0.01824652,0.00077725016,0.0022343097],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98832345,0.0017538224,0.0021709278,0.0021572774,0.004922818,0.00067174324],"domain_scores_gemma":[0.9253538,0.019277172,0.0147054875,0.0071619353,0.030556628,0.0029449426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007974399,0.003386478,0.002537016,0.031370003,0.0022283474,0.008640285,0.0042938674,0.003208532,0.0032320207],"category_scores_gemma":[0.055108044,0.0010734305,0.0022912193,0.013894321,0.0012028625,0.009574802,0.004329927,0.0039388873,0.00430967],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006020185,0.0012476302,0.18819557,0.0020495818,0.0002596628,0.0011025207,0.00452514,0.008634717,0.013580451,0.031031476,0.05343221,0.695339],"study_design_scores_gemma":[0.00039669414,0.0013377183,0.21619993,0.0038121317,0.00072699273,0.0035932136,0.008960744,0.45112625,0.019646905,0.12499855,0.16835646,0.0008444483],"about_ca_topic_score_codex":0.011508936,"about_ca_topic_score_gemma":0.00791329,"teacher_disagreement_score":0.031370003,"about_ca_system_score_codex":0.002401435,"about_ca_system_score_gemma":0.0039537037,"threshold_uncertainty_score":0.042173147},"labels":[],"label_agreement":null},{"id":"W2883564831","doi":"10.1145/3196398.3196409","title":"A design structure matrix approach for measuring co-change-modularity of software products","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada; Western Canada Research Grid; Compute Canada","keywords":"Design structure matrix; Computer science; Modularity (biology); Metric (unit); Software system; Theoretical computer science; Data mining; Software metric; Software; Change impact analysis; Representation (politics); Programming language; Software construction; Mathematics","score_opus":0.09623937090276362,"score_gpt":0.3114389159828748,"score_spread":0.2151995450801112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883564831","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.088296205,0.00028764896,0.9056476,0.00016388431,0.00003305615,0.00039217077,0.0005301554,0.0006148681,0.004034321],"genre_scores_gemma":[0.4253281,0.00011623209,0.5718797,0.000045492143,0.000016549044,0.0005744566,0.000686982,0.00010046183,0.0012521351],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956399,0.001040561,0.00034616832,0.0005323449,0.0023080506,0.00013288033],"domain_scores_gemma":[0.9859201,0.0060360455,0.0032905156,0.001474098,0.0029589785,0.00032031417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029515561,0.0010772793,0.0004934095,0.0076684547,0.0005728806,0.0012865358,0.00077802816,0.0009052654,0.0021670887],"category_scores_gemma":[0.017928021,0.00043896938,0.000855557,0.0057717357,0.0009566621,0.0022912878,0.0012133336,0.00094915606,0.0003676484],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000330311,0.0006614576,0.057382647,0.0009507606,0.0006269522,0.0002761234,0.002109607,0.1848608,0.057237647,0.13217813,0.004331077,0.55905455],"study_design_scores_gemma":[0.000082780556,0.0010529627,0.04613654,0.0001108705,0.00022968692,0.0008054942,0.00079934165,0.8060834,0.029247858,0.09732024,0.017947948,0.00018286951],"about_ca_topic_score_codex":0.0028696952,"about_ca_topic_score_gemma":0.004447123,"teacher_disagreement_score":0.0076684547,"about_ca_system_score_codex":0.0021085336,"about_ca_system_score_gemma":0.001047398,"threshold_uncertainty_score":0.015609562},"labels":[],"label_agreement":null},{"id":"W2883986603","doi":"10.1145/3196398.3196438","title":"CLEVER","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Ubisoft (Canada)","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Commit; Computer science; Leverage (statistics); Reuse; Process (computing); Metric (unit); Software; Software maintenance; Software engineering; Software bug; Field (mathematics); Software development; Artificial intelligence; Database; Programming language; Engineering","score_opus":0.02048636646562555,"score_gpt":0.276926340176175,"score_spread":0.25643997371054944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883986603","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023820955,0.0041198735,0.40605763,0.005619027,0.0038915873,0.002089881,0.04327318,0.30832615,0.20280175],"genre_scores_gemma":[0.16545826,0.0017296219,0.49965012,0.0060428376,0.00079043984,0.001376716,0.13866012,0.043589372,0.14270246],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9930922,0.0017446902,0.00050418417,0.0019182015,0.0021614418,0.0005791696],"domain_scores_gemma":[0.9868433,0.0054906034,0.00058856356,0.004441601,0.002201234,0.0004347517],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0039159227,0.0025729707,0.0015847039,0.0040109335,0.0015810146,0.0040693334,0.003609729,0.0036261738,0.11505727],"category_scores_gemma":[0.025452571,0.0011116888,0.001728533,0.0026876172,0.00095887727,0.005045115,0.0034805718,0.0027476482,0.08731185],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046811145,0.00023940303,0.0029776108,0.0010747099,0.00011491899,0.00028803613,0.0002390175,0.006769877,0.0025713914,0.015512323,0.6362502,0.3334944],"study_design_scores_gemma":[0.0003284152,0.00024115117,0.002509158,0.00032747866,0.00008496228,0.001255864,0.00031128913,0.10911124,0.011861154,0.036034774,0.8377776,0.00015685394],"about_ca_topic_score_codex":0.003282312,"about_ca_topic_score_gemma":0.0069285356,"teacher_disagreement_score":0.8849427,"about_ca_system_score_codex":0.0011712952,"about_ca_system_score_gemma":0.0024983375,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2884048294","doi":"10.1145/3196398.3196435","title":"Studying the relationship between exception handling practices and post-release defects","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Exception handling; Java; Computer science; Software quality; Software; Quality (philosophy); Reliability (semiconductor); Software bug; Work flow; Software release life cycle; Software engineering; Software development; Programming language; Power (physics); Engineering; Industrial engineering","score_opus":0.1030517818798358,"score_gpt":0.34820214274140243,"score_spread":0.24515036086156664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884048294","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.991231,0.000277054,0.0072633405,0.00023384973,0.000012064727,0.000035764508,0.00021119231,0.00006801496,0.00066762854],"genre_scores_gemma":[0.9976775,0.00011570909,0.0015724426,0.000021967871,0.000012887253,0.000026304717,0.00028618806,0.000021679558,0.0002653712],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9878593,0.003803821,0.001338878,0.0024794792,0.003448873,0.0010695957],"domain_scores_gemma":[0.50867844,0.33156466,0.12045468,0.014330138,0.020663677,0.004308472],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014071492,0.0008051298,0.000512506,0.0034191017,0.0004804069,0.0022491298,0.0016510385,0.0014305688,0.0018411374],"category_scores_gemma":[0.12326156,0.0006441038,0.0014297162,0.002955076,0.0011037871,0.0037541706,0.0014743353,0.002406549,0.000465817],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000052677256,0.00011984458,0.9885346,0.00005842368,0.00024259098,0.00013656341,0.00046376212,0.0036677616,0.00026093487,0.0003215666,0.000094575764,0.006046704],"study_design_scores_gemma":[0.000006394425,0.0003838086,0.97341925,0.00006105063,0.00012652035,0.00025159007,0.00084226736,0.02293002,0.00050230447,0.0010255707,0.0004170025,0.000034256947],"about_ca_topic_score_codex":0.006253882,"about_ca_topic_score_gemma":0.008037262,"teacher_disagreement_score":0.014071492,"about_ca_system_score_codex":0.0011246669,"about_ca_system_score_gemma":0.001876557,"threshold_uncertainty_score":0.07441807},"labels":[],"label_agreement":null},{"id":"W2884383315","doi":"10.1145/3195546.3206423","title":"A project on software defect prevention at commit-time","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ubisoft (Canada); Concordia University","funders":"Concordia University","keywords":"Commit; Computer science; Novelty; Coding (social sciences); Code review; Software; Software engineering; Software quality; Code (set theory); Software development; Programming language; Database","score_opus":0.026633061175849,"score_gpt":0.30003179591226736,"score_spread":0.27339873473641835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884383315","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18781951,0.004514365,0.67832345,0.007817885,0.0030779128,0.0067649577,0.005263848,0.04776712,0.058650915],"genre_scores_gemma":[0.22055389,0.0015437916,0.67425877,0.0024174459,0.0003183264,0.002121379,0.013575168,0.005706996,0.079504244],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9827571,0.0067761773,0.0005150059,0.0027531008,0.005710692,0.0014879369],"domain_scores_gemma":[0.95947516,0.011713375,0.0011156827,0.009242859,0.012935298,0.0055176495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015709968,0.0015069261,0.0013067382,0.003850214,0.0022972964,0.0040512104,0.002752352,0.0031647603,0.0129753025],"category_scores_gemma":[0.025012359,0.0010715481,0.0014834802,0.001807522,0.0016140299,0.0041357903,0.004523503,0.0032548402,0.0040898346],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017719404,0.0029356277,0.0081507815,0.000989459,0.00025000115,0.00074509537,0.0031524166,0.0127673,0.044491533,0.022799604,0.06963394,0.83231235],"study_design_scores_gemma":[0.0026964787,0.016422225,0.04053344,0.0013350601,0.0006013669,0.0036565003,0.003966131,0.12508577,0.1587235,0.02509793,0.6213396,0.0005419768],"about_ca_topic_score_codex":0.006418166,"about_ca_topic_score_gemma":0.0067020217,"teacher_disagreement_score":0.015709968,"about_ca_system_score_codex":0.0026360631,"about_ca_system_score_gemma":0.008405829,"threshold_uncertainty_score":0.08308321},"labels":[],"label_agreement":null},{"id":"W2884390033","doi":"10.1145/3196398.3196436","title":"Large-scale analysis of the co-commit patterns of the active developers in github's top repositories","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Commit; Computer science; Reputation; World Wide Web; Ranking (information retrieval); Scale (ratio); Data science; Database; Information retrieval","score_opus":0.010709485535518648,"score_gpt":0.26943728332676725,"score_spread":0.2587277977912486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884390033","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99483705,0.00031206178,0.0015074939,0.00013196828,0.000012756494,0.000036279893,0.0018613512,0.00042538947,0.0008755617],"genre_scores_gemma":[0.98556006,0.00022963168,0.0038406793,0.00003271457,0.000026824486,0.00008432738,0.008786914,0.00018133668,0.0012575376],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99730515,0.0005502331,0.00015536297,0.0005473196,0.00103719,0.00040472596],"domain_scores_gemma":[0.9739043,0.009026269,0.0057231174,0.0036494145,0.005533971,0.0021630037],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002124296,0.0005749185,0.0004986573,0.00992082,0.0010106667,0.0013930093,0.0010669886,0.00062477123,0.0006658254],"category_scores_gemma":[0.018513298,0.0003988804,0.0005293773,0.00993632,0.0009001298,0.0021330968,0.0022638277,0.0009902088,0.00051592843],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039337963,0.0004213587,0.8803577,0.00045068993,0.0004494452,0.0015746752,0.00689699,0.011895515,0.0093462365,0.001855188,0.013975565,0.0723833],"study_design_scores_gemma":[0.000024573534,0.00013016978,0.9477544,0.000048663624,0.00008611355,0.00070189347,0.002880707,0.037109617,0.0037787284,0.0009269891,0.006494086,0.00006418222],"about_ca_topic_score_codex":0.0141018685,"about_ca_topic_score_gemma":0.025380917,"teacher_disagreement_score":0.0141018685,"about_ca_system_score_codex":0.0009593259,"about_ca_system_score_gemma":0.0008230508,"threshold_uncertainty_score":0.028039575},"labels":[],"label_agreement":null},{"id":"W2884493935","doi":"10.1145/3196398.3196463","title":"Studying developer build issues and debugger usage via timeline analysis in visual studio IDE","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Timeline; Debugger; Computer science; Debugging; Software engineering; Plug-in; Microsoft Visual Studio; Session (web analytics); Source code; Workflow; Software; World Wide Web; Programming language; Database","score_opus":0.021489812266390816,"score_gpt":0.3301445447793923,"score_spread":0.3086547325130015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884493935","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9808623,0.00041366235,0.009167266,0.00010878338,0.000020304686,0.00006280697,0.005930986,0.0012280238,0.0022058603],"genre_scores_gemma":[0.9257235,0.00047600234,0.03957235,0.000057235346,0.000042287786,0.00031294575,0.03066541,0.00074165117,0.002408695],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9962446,0.0010127571,0.00042317115,0.00079532247,0.0012541424,0.00027001637],"domain_scores_gemma":[0.9608825,0.024421198,0.006855966,0.0026166625,0.004004576,0.0012191295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042334264,0.000504795,0.0004475948,0.0070875743,0.00038186082,0.0012981651,0.00072267937,0.00050607795,0.0005676903],"category_scores_gemma":[0.021648917,0.00037289425,0.000409482,0.005209254,0.0003768909,0.0016700092,0.001037174,0.00092896976,0.0006085476],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001394837,0.0009043984,0.6097682,0.0013120217,0.00028487484,0.00091254234,0.02140836,0.010247537,0.025684081,0.0024838198,0.023184335,0.30241498],"study_design_scores_gemma":[0.000106935186,0.0007931992,0.8627885,0.0002479723,0.00014962739,0.0008966772,0.0063655744,0.06033729,0.018832954,0.0023349947,0.046958238,0.00018813695],"about_ca_topic_score_codex":0.0039844927,"about_ca_topic_score_gemma":0.0077625746,"teacher_disagreement_score":0.0070875743,"about_ca_system_score_codex":0.000601665,"about_ca_system_score_gemma":0.0007612747,"threshold_uncertainty_score":0.022388756},"labels":[],"label_agreement":null},{"id":"W2884885238","doi":"10.1145/3196398.3196420","title":"Exploring the use of automated API migrating techniques in practice","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Documentation; Computer science; Android (operating system); Application programming interface; Software documentation; World Wide Web; Software; Software engineering; Data science; Software development; Programming language; Software development process; Operating system","score_opus":0.16086013620484027,"score_gpt":0.3423152367888403,"score_spread":0.18145510058400002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884885238","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34253818,0.0021554155,0.61222863,0.0058088894,0.00019124943,0.0010598915,0.00016809195,0.01275458,0.023095066],"genre_scores_gemma":[0.424583,0.0008034652,0.5703976,0.00033975852,0.00003143837,0.00029573354,0.00025180273,0.00085916906,0.0024381073],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.977384,0.013182915,0.0011765547,0.0025030987,0.0046324558,0.0011210479],"domain_scores_gemma":[0.91750795,0.047810115,0.0071371095,0.01708018,0.008821997,0.0016426917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019341856,0.0014177022,0.0007026979,0.0030188004,0.0024446284,0.005498367,0.004750274,0.0028580474,0.0030959875],"category_scores_gemma":[0.08913697,0.001510351,0.00077328284,0.0018946093,0.0025331238,0.008815524,0.004652846,0.0030582193,0.0018393095],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055714505,0.0017398542,0.024212677,0.001975794,0.00012939193,0.0015675885,0.035886936,0.018251173,0.027343407,0.018248625,0.0071419408,0.8629455],"study_design_scores_gemma":[0.00064154266,0.0036109565,0.028933704,0.003866834,0.0005221807,0.00699438,0.058568817,0.45422637,0.063454516,0.085310936,0.2930623,0.0008075484],"about_ca_topic_score_codex":0.0034678369,"about_ca_topic_score_gemma":0.006387709,"teacher_disagreement_score":0.019341856,"about_ca_system_score_codex":0.0020690006,"about_ca_system_score_gemma":0.0049145874,"threshold_uncertainty_score":0.10229075},"labels":[],"label_agreement":null},{"id":"W2884905421","doi":"10.1145/3196398.3196421","title":"Studying the impact of adopting continuous integration on the delivery time of pull requests","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Metric (unit); Software; Software engineering; Software development; Key (lock); Work (physics); Operating system; Operations management; Engineering","score_opus":0.03053442553353595,"score_gpt":0.2945840618263738,"score_spread":0.26404963629283784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884905421","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99333763,0.00042645712,0.003942711,0.00028442324,0.000029741437,0.000049902053,0.0004846568,0.00022775895,0.001216755],"genre_scores_gemma":[0.99635303,0.00014363883,0.0018959094,0.000045881534,0.000021816239,0.000045296863,0.0008689718,0.000080671794,0.0005448469],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98841864,0.004336159,0.0009370853,0.002091304,0.002527504,0.0016893051],"domain_scores_gemma":[0.6270815,0.30243775,0.041753143,0.010630393,0.01127533,0.006821838],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02211941,0.0012373622,0.00069689483,0.0024436493,0.0006980321,0.0032995434,0.0014651687,0.001488904,0.003329666],"category_scores_gemma":[0.16326378,0.00074192433,0.0011653218,0.00280964,0.001114089,0.0041225897,0.0018219833,0.0044192425,0.0013753322],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015571355,0.0011608965,0.85844284,0.0004938188,0.0007141644,0.0005599469,0.0018008407,0.07779469,0.005015367,0.0017596749,0.0020638662,0.048636712],"study_design_scores_gemma":[0.00012908826,0.002787856,0.726617,0.00012510554,0.0005456455,0.00059197407,0.0023360737,0.2565502,0.005042386,0.0020272627,0.0030926059,0.0001548055],"about_ca_topic_score_codex":0.009098263,"about_ca_topic_score_gemma":0.0071568806,"teacher_disagreement_score":0.02211941,"about_ca_system_score_codex":0.0014933073,"about_ca_system_score_gemma":0.0019325762,"threshold_uncertainty_score":0.116980076},"labels":[],"label_agreement":null},{"id":"W2885308680","doi":"10.1109/tse.2018.2861735","title":"Leveraging Historical Associations between Requirements and Source Code to Identify Impacted Classes","year":2018,"lang":"en","type":"preprint","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"National Science Foundation","keywords":"Computer science; Set (abstract data type); Class (philosophy); Code smell; Intuition; Similarity (geometry); Semantic similarity; Source code; Data mining; Locality; Software; Information retrieval; Artificial intelligence; Software development; Software quality; Programming language","score_opus":0.05722962105556756,"score_gpt":0.32059560617191873,"score_spread":0.2633659851163512,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885308680","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8881859,0.00078379764,0.100551955,0.00028190383,0.000049741142,0.00018836885,0.003072862,0.0029020698,0.0039834147],"genre_scores_gemma":[0.9374622,0.00024052252,0.055643473,0.000049383303,0.000048406535,0.00010284014,0.0054877084,0.00018450122,0.0007809915],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967895,0.00053772004,0.00034646937,0.0007093723,0.0014841422,0.00013273695],"domain_scores_gemma":[0.9666235,0.015117707,0.009384036,0.0025937757,0.0055270013,0.00075395685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022121156,0.00077404885,0.00052312156,0.008377003,0.00046762524,0.0010737951,0.0006206077,0.0006974104,0.0007181677],"category_scores_gemma":[0.026320703,0.0003050338,0.0007368044,0.0035869158,0.00038141332,0.0028402493,0.0009373887,0.0008100573,0.00055506954],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005011697,0.00044993425,0.5991451,0.0005187772,0.00032981558,0.00074597273,0.0011369151,0.034487575,0.019286083,0.0011120723,0.0035619927,0.3387246],"study_design_scores_gemma":[0.000031176416,0.00076773233,0.48953658,0.00010161412,0.00017297898,0.0010604905,0.0006884826,0.48340124,0.015244018,0.0027843437,0.0060953237,0.000115982126],"about_ca_topic_score_codex":0.0043524876,"about_ca_topic_score_gemma":0.010159703,"teacher_disagreement_score":0.008377003,"about_ca_system_score_codex":0.0006277714,"about_ca_system_score_gemma":0.00075364363,"threshold_uncertainty_score":0.011698961},"labels":[],"label_agreement":null},{"id":"W2885541650","doi":"10.1145/3180155.3182514","title":"Are fix-inducing changes a moving target?","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Granularity; Computer science; Process (computing); Code (set theory); Reliability engineering; Programming language; Engineering","score_opus":0.03183609157711317,"score_gpt":0.27563106616453126,"score_spread":0.24379497458741808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885541650","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34272802,0.008467666,0.5579317,0.018168565,0.00218697,0.00042227557,0.002744883,0.010711373,0.056638554],"genre_scores_gemma":[0.9081762,0.0017894204,0.06975389,0.0029151735,0.00053196517,0.0002129437,0.0016817248,0.0018543927,0.013084291],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9948945,0.00068179285,0.00023448313,0.0016937148,0.002045676,0.0004498874],"domain_scores_gemma":[0.9552275,0.020766733,0.0073449523,0.00919244,0.006287233,0.001181217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038017898,0.0009931518,0.00095466216,0.0016583257,0.0011774601,0.0028896593,0.0020802903,0.0035729068,0.014713646],"category_scores_gemma":[0.05374302,0.0009560475,0.00083054317,0.001179806,0.0017640499,0.008060005,0.001985733,0.0027862513,0.006186185],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010447945,0.00060937816,0.10998867,0.0015454893,0.0002580235,0.002345837,0.002082278,0.011577118,0.03249126,0.09873686,0.04024418,0.6990762],"study_design_scores_gemma":[0.00025223434,0.001538624,0.17646208,0.0011647919,0.0005087585,0.009969789,0.0049211485,0.12638414,0.08135914,0.40211415,0.19493882,0.00038631805],"about_ca_topic_score_codex":0.0017781113,"about_ca_topic_score_gemma":0.0017946478,"teacher_disagreement_score":0.014713646,"about_ca_system_score_codex":0.0011023406,"about_ca_system_score_gemma":0.0011738553,"threshold_uncertainty_score":0.049222052},"labels":[],"label_agreement":null},{"id":"W2885871161","doi":"10.1145/3180155.3182518","title":"Analyzing a decade of Linux system calls","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; System call; Linux kernel; Operating system; Kernel (algebra); sysfs; Application programming interface; Process (computing); Interface (matter); Software engineering","score_opus":0.01915276403524373,"score_gpt":0.28215258503460355,"score_spread":0.2629998209993598,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885871161","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99434495,0.0007011655,0.00028265556,0.0003669662,0.000015497917,0.000008483173,0.00082118023,0.000028179764,0.0034310066],"genre_scores_gemma":[0.99737835,0.00039599347,0.00028409515,0.000089342604,0.000032785876,0.000010045973,0.0011770971,0.000027038228,0.00060516375],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99724126,0.00047000588,0.00030932223,0.00044558672,0.001199138,0.00033461643],"domain_scores_gemma":[0.95996886,0.014795794,0.014207274,0.0015723754,0.0080606835,0.0013950575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023721692,0.00014693468,0.00019579582,0.004141403,0.0006806568,0.0019402159,0.0005812459,0.0006535542,0.0011722762],"category_scores_gemma":[0.032156963,0.00026175572,0.00017970095,0.0059784623,0.0007651592,0.0020518436,0.0008809201,0.001158744,0.0005203729],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000247271,0.00013819865,0.9317564,0.00011461604,0.000046913166,0.00023678168,0.010870825,0.0007726737,0.0011694104,0.0022500362,0.004419353,0.04797757],"study_design_scores_gemma":[0.0000032363255,0.0000627312,0.9826649,0.00005678263,0.000014771231,0.00016328688,0.0045828754,0.0011890419,0.0004441833,0.0002924103,0.01050552,0.000020253307],"about_ca_topic_score_codex":0.018302005,"about_ca_topic_score_gemma":0.019130504,"teacher_disagreement_score":0.018302005,"about_ca_system_score_codex":0.0018871354,"about_ca_system_score_gemma":0.00094882504,"threshold_uncertainty_score":0.0363909},"labels":[],"label_agreement":null},{"id":"W2886024027","doi":"10.1109/qrs.2018.00048","title":"The State of Practice on Virtual Reality (VR) Applications: An Exploratory Study on Github and Stack Overflow","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Virtual reality; JavaScript; Software; Exploratory research; Perspective (graphical); Software development; Multimedia; State (computer science); World Wide Web; Human–computer interaction; Software engineering; Artificial intelligence","score_opus":0.043432985563012165,"score_gpt":0.34560851523144975,"score_spread":0.3021755296684376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886024027","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99720675,0.00015967093,0.0008255496,0.000271972,0.00000322166,0.00003856953,0.000034196368,0.000019507464,0.001440545],"genre_scores_gemma":[0.99705625,0.00030640076,0.0016604009,0.00012427913,0.000007045319,0.000099014636,0.00006765451,0.00005379586,0.00062524533],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9814337,0.011157781,0.0009283908,0.001448682,0.003734405,0.0012971142],"domain_scores_gemma":[0.93395686,0.046886317,0.009589626,0.0018673309,0.0050896253,0.0026102755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013222449,0.00045071795,0.0005737908,0.005827055,0.002435823,0.004180314,0.0013672705,0.0012050833,0.0011442978],"category_scores_gemma":[0.0657575,0.00056960894,0.00026600776,0.0040239454,0.0058234506,0.006206816,0.0043595103,0.0011403782,0.00026136194],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006025222,0.00018361148,0.049187843,0.00033223146,0.000017803915,0.000956173,0.9144745,0.0001170493,0.0020893824,0.0022789699,0.0008678417,0.029434348],"study_design_scores_gemma":[0.000011825945,0.0002899336,0.12418513,0.00055477966,0.00001986858,0.001036987,0.85095036,0.00095163257,0.001074984,0.0007946352,0.020052101,0.00007773331],"about_ca_topic_score_codex":0.005062109,"about_ca_topic_score_gemma":0.008129125,"teacher_disagreement_score":0.013222449,"about_ca_system_score_codex":0.00314606,"about_ca_system_score_gemma":0.002369403,"threshold_uncertainty_score":0.06992787},"labels":[],"label_agreement":null},{"id":"W2886117034","doi":"10.1145/3236024.3236065","title":"Improving IR-based bug localization with context-aware query reformulation","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":106,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Information retrieval; Context (archaeology); Query expansion; Query language; Baseline (sea); State (computer science); Data mining; Database; Programming language","score_opus":0.012833199998753323,"score_gpt":0.24407213111559414,"score_spread":0.23123893111684082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886117034","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.090785615,0.0065038013,0.82780695,0.0012637392,0.0002845696,0.00063995016,0.001239629,0.06814893,0.0033268302],"genre_scores_gemma":[0.29476798,0.0015503869,0.69401544,0.00078024546,0.00034591302,0.00023724855,0.0033458893,0.0014672835,0.0034896347],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99528223,0.0012879162,0.00046718965,0.0010296989,0.0016377791,0.0002950518],"domain_scores_gemma":[0.9897746,0.0043812543,0.0012444536,0.0019523933,0.0024541242,0.0001932145],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031837486,0.0024732535,0.002542904,0.006000862,0.00083050516,0.0017054123,0.002272656,0.0015241352,0.0032583321],"category_scores_gemma":[0.01679026,0.0006376089,0.0018752553,0.0035122891,0.0009161893,0.004312466,0.0020384805,0.0016555901,0.0027987757],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052653113,0.00047368984,0.005036344,0.0013958813,0.00017724144,0.0005395287,0.0012478083,0.014190307,0.1126029,0.0033365623,0.025360163,0.83511305],"study_design_scores_gemma":[0.0004574257,0.0016631639,0.010126347,0.00020733582,0.0012046597,0.0032583608,0.0014675326,0.7129163,0.21285188,0.010706581,0.044787444,0.00035297146],"about_ca_topic_score_codex":0.007998926,"about_ca_topic_score_gemma":0.00573896,"teacher_disagreement_score":0.007998926,"about_ca_system_score_codex":0.0010695521,"about_ca_system_score_gemma":0.0020178354,"threshold_uncertainty_score":0.016837478},"labels":[],"label_agreement":null},{"id":"W2886347592","doi":"10.1007/s10664-018-9636-3","title":"An empirical study on the issue reports with questions raised during the issue resolving process","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Beihang University","keywords":"Computer science; Process (computing); Data science; Eclipse; Empirical research","score_opus":0.020315733376175562,"score_gpt":0.3259201812352616,"score_spread":0.3056044478590861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886347592","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99540436,0.00015465186,0.0013047103,0.00018165063,0.00002406068,0.0002180242,0.00014011528,0.000021034388,0.0025513628],"genre_scores_gemma":[0.9942584,0.0002095655,0.0026597378,0.00020020522,0.000058208563,0.00027718156,0.0004603054,0.000031715248,0.0018446964],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9689794,0.022005502,0.0022598724,0.0014708978,0.004544881,0.0007393977],"domain_scores_gemma":[0.40338665,0.4972746,0.05679219,0.016714202,0.021221,0.0046113343],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.021783276,0.0005484095,0.00046256144,0.0031251279,0.0021017182,0.0034583104,0.0015079395,0.001776516,0.004710596],"category_scores_gemma":[0.3017368,0.00063842064,0.00050533947,0.0034554426,0.0019474891,0.0035363801,0.0023301325,0.003128713,0.0013211583],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032394512,0.025179578,0.70758355,0.0012213668,0.00022906666,0.0021650908,0.12995586,0.0010297414,0.007862128,0.0046789423,0.0036278402,0.113227464],"study_design_scores_gemma":[0.00041423336,0.0080002025,0.8304944,0.00052365306,0.0003618708,0.0027380919,0.114890926,0.0059742974,0.012439423,0.0025858632,0.021338888,0.00023806165],"about_ca_topic_score_codex":0.0021330796,"about_ca_topic_score_gemma":0.0020130398,"teacher_disagreement_score":0.9782167,"about_ca_system_score_codex":0.0011731879,"about_ca_system_score_gemma":0.0017594008,"threshold_uncertainty_score":0.11520237},"labels":[],"label_agreement":null},{"id":"W2886769913","doi":"10.1016/j.jss.2018.08.032","title":"Improving reusability of software libraries through usage pattern mining","year":2018,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":68,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; École de Technologie Supérieure; Université de Montréal; Concordia University","funders":"","keywords":"Computer science; Reusability; Reuse; World Wide Web; Software; Software engineering; Set (abstract data type); Database; Operating system; Engineering","score_opus":0.024470703776739766,"score_gpt":0.25866016959175675,"score_spread":0.23418946581501698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886769913","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8088551,0.0025723493,0.17379417,0.0005542971,0.0000695912,0.00022781285,0.0025188562,0.008339624,0.003068181],"genre_scores_gemma":[0.8585246,0.0007221866,0.13260916,0.0000967665,0.000048346697,0.00012553821,0.005642263,0.0005654855,0.0016656246],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99549353,0.0007389984,0.00074283953,0.00078861834,0.0019364076,0.00029961378],"domain_scores_gemma":[0.9797736,0.008752619,0.0034227974,0.0034736546,0.004195243,0.00038204333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024698102,0.0010083776,0.0010921968,0.010083488,0.00064803095,0.0018017568,0.0017110915,0.00088946073,0.0007457759],"category_scores_gemma":[0.019631358,0.00054583146,0.0016330524,0.0059796274,0.00038035,0.0039728303,0.0012990155,0.00097305694,0.00056217465],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004181254,0.0009844345,0.12549077,0.0007732716,0.0005273476,0.00048459164,0.000569745,0.017562745,0.01995501,0.0012293261,0.003015639,0.82898897],"study_design_scores_gemma":[0.00010997935,0.00096218294,0.0978207,0.00035361067,0.0012197667,0.001582885,0.0008914081,0.82060176,0.053946145,0.012476105,0.009896528,0.00013894383],"about_ca_topic_score_codex":0.0045557613,"about_ca_topic_score_gemma":0.00887349,"teacher_disagreement_score":0.010083488,"about_ca_system_score_codex":0.0005078485,"about_ca_system_score_gemma":0.0015648811,"threshold_uncertainty_score":0.013061762},"labels":[],"label_agreement":null},{"id":"W2886876310","doi":"10.1002/j.2334-5837.2018.00501.x","title":"Delivering Better Projects on Time by Ensuring Requirements Quality Upfront","year":2018,"lang":"en","type":"article","venue":"INCOSE International Symposium","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Technical University of Nova Scotia","funders":"","keywords":"Timeline; Computer science; Criticality; Risk analysis (engineering); Automation; Quality (philosophy); Domain (mathematical analysis); Process management; Control (management); Systems engineering; Engineering; Business","score_opus":0.027531673562654495,"score_gpt":0.3110211638853228,"score_spread":0.2834894903226683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886876310","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19838503,0.0007524113,0.70074534,0.017792704,0.00025010505,0.0011710656,0.00030101262,0.0041206987,0.07648158],"genre_scores_gemma":[0.619747,0.00043297687,0.37048063,0.00078390015,0.000083607,0.00039963354,0.00035138417,0.00094049,0.006780433],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9554316,0.020903578,0.0028963555,0.0021066472,0.01644272,0.0022190826],"domain_scores_gemma":[0.8173134,0.06533015,0.029728122,0.033508867,0.04868175,0.0054376977],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02857281,0.0009360153,0.000515862,0.0029253839,0.0018068469,0.008389057,0.0019169196,0.0016256549,0.0063249767],"category_scores_gemma":[0.115336955,0.0007983407,0.0005722866,0.0015646844,0.0017158819,0.0067502884,0.0044771004,0.00219525,0.0032317918],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033312163,0.0009867668,0.03397255,0.001058881,0.00012146685,0.00085756194,0.010526716,0.04124616,0.044933923,0.10304726,0.025117965,0.7377976],"study_design_scores_gemma":[0.00030219826,0.00177907,0.07443192,0.0032578094,0.00021575048,0.002840156,0.022134192,0.186691,0.0907152,0.22584549,0.3913533,0.0004339214],"about_ca_topic_score_codex":0.0035111767,"about_ca_topic_score_gemma":0.0038100265,"teacher_disagreement_score":0.02857281,"about_ca_system_score_codex":0.003133111,"about_ca_system_score_gemma":0.010802321,"threshold_uncertainty_score":0.15110928},"labels":[],"label_agreement":null},{"id":"W2886906914","doi":"10.1145/3180155.3182538","title":"Measuring program comprehension","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Program comprehension; Comprehension; Computer science; Software maintenance; Software engineering; Software; Software development; Programming language; Software system","score_opus":0.06453357335296482,"score_gpt":0.2990361902343799,"score_spread":0.23450261688141505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886906914","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9429964,0.001349485,0.025009884,0.00023271701,0.00008790579,0.0011312971,0.002335725,0.0014447696,0.025411801],"genre_scores_gemma":[0.95663863,0.00087740755,0.03142331,0.00022699467,0.000054207623,0.0013411035,0.004625715,0.00017584974,0.0046368316],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9910972,0.0029866768,0.0010969738,0.0011303468,0.0033008964,0.00038782554],"domain_scores_gemma":[0.946309,0.030094633,0.008408542,0.002789627,0.010663456,0.0017346943],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.005584523,0.0008732699,0.0006240945,0.0035601587,0.00047161675,0.0012655696,0.0007774782,0.00094603316,0.004007666],"category_scores_gemma":[0.054844324,0.0003261331,0.0007491597,0.0018394662,0.00045962783,0.0021767903,0.0013385065,0.0012198711,0.0014302691],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011767271,0.002290486,0.4524032,0.0027562676,0.0007613449,0.00022352078,0.02066644,0.001836069,0.041512184,0.0024593314,0.011821385,0.462093],"study_design_scores_gemma":[0.00011961934,0.003280879,0.9463077,0.00036226033,0.00036752457,0.00066142314,0.0034806286,0.004629473,0.019382337,0.0025438936,0.018720085,0.00014424417],"about_ca_topic_score_codex":0.00092143187,"about_ca_topic_score_gemma":0.0010898608,"teacher_disagreement_score":0.99441546,"about_ca_system_score_codex":0.00059433206,"about_ca_system_score_gemma":0.0007367718,"threshold_uncertainty_score":0.029534161},"labels":[],"label_agreement":null},{"id":"W2886935593","doi":"","title":"Using Natural Language Processing for Documentation Assist.","year":2018,"lang":"en","type":"article","venue":"National Conference on Artificial Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Documentation; Computer science; Internal documentation; World Wide Web; Technical documentation; Cursor (databases); Natural language; Software engineering; Database; Programming language; Software; Natural language processing; Software development","score_opus":0.17915399153434525,"score_gpt":0.43928851831647003,"score_spread":0.2601345267821248,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886935593","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042557027,0.0007503328,0.9522797,0.0012733542,0.00021548133,0.00036584516,0.00291986,0.027334988,0.010604634],"genre_scores_gemma":[0.030356498,0.0005811025,0.9558381,0.00040866178,0.00008655839,0.00030634864,0.005972871,0.0017802963,0.0046696104],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99491334,0.0021025233,0.00055429916,0.0008285169,0.0014875252,0.00011369841],"domain_scores_gemma":[0.982624,0.011260488,0.0015024713,0.002087756,0.0022717707,0.00025340208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042866934,0.001614292,0.00056054594,0.0045338627,0.001113058,0.0049885404,0.0017530926,0.0016215846,0.0110930335],"category_scores_gemma":[0.018530857,0.0007870942,0.001659582,0.002440411,0.0014960947,0.006344643,0.0024483162,0.0021375786,0.009647251],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035187238,0.0002999317,0.0020155627,0.0035890865,0.0001574654,0.0017612607,0.0048599183,0.0061181434,0.037576016,0.091636226,0.09037598,0.7612585],"study_design_scores_gemma":[0.00014085563,0.00018470433,0.0018534928,0.0014589635,0.00014538398,0.0032468678,0.002492194,0.11221711,0.046658516,0.21322064,0.6180971,0.00028413069],"about_ca_topic_score_codex":0.0023008303,"about_ca_topic_score_gemma":0.0034233646,"teacher_disagreement_score":0.0110930335,"about_ca_system_score_codex":0.0011799257,"about_ca_system_score_gemma":0.002283242,"threshold_uncertainty_score":0.03710985},"labels":[],"label_agreement":null},{"id":"W2887081063","doi":"10.1145/3205455.3205622","title":"Towards the automated recovery of complex temporal API-usage patterns","year":2018,"lang":"en","type":"article","venue":"Proceedings of the Genetic and Evolutionary Computation Conference","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Concordia University","funders":"","keywords":"Computer science","score_opus":0.035859079483935574,"score_gpt":0.2650376949175621,"score_spread":0.22917861543362655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2887081063","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037164703,0.00010851878,0.9548517,0.00032224978,0.00002342179,0.00012097789,0.00028734095,0.0059605814,0.0011605823],"genre_scores_gemma":[0.15416014,0.00013085664,0.84162205,0.00015005049,0.000015725644,0.0001272037,0.0010920264,0.0010853742,0.0016165876],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99759454,0.0005399918,0.00016364532,0.00060594385,0.00088710914,0.00020865623],"domain_scores_gemma":[0.99248475,0.003077346,0.0011593957,0.001955972,0.0011671135,0.0001554527],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020195693,0.0013236016,0.00074212236,0.0020576154,0.0008089808,0.0016466767,0.002135145,0.0016723593,0.001398607],"category_scores_gemma":[0.012446669,0.0009390662,0.001725139,0.0018133118,0.0012918076,0.0019234519,0.0021052323,0.0025650165,0.00077684096],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002711437,0.00060592004,0.025462238,0.00065543514,0.00025273155,0.0011979108,0.0014298554,0.22531395,0.046238624,0.027535256,0.009858168,0.6611787],"study_design_scores_gemma":[0.000026886119,0.000038941485,0.0014106887,0.0000509108,0.00005072267,0.0003944407,0.00018314757,0.9520376,0.018036412,0.023015734,0.0047241114,0.000030515484],"about_ca_topic_score_codex":0.0064279535,"about_ca_topic_score_gemma":0.0066098985,"teacher_disagreement_score":0.0064279535,"about_ca_system_score_codex":0.0008241605,"about_ca_system_score_gemma":0.002781487,"threshold_uncertainty_score":0.012781084},"labels":[],"label_agreement":null},{"id":"W2887758210","doi":"10.1145/3194104.3194110","title":"A replication study","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Preprocessor; Ensemble learning; Random forest; Replication (statistics); Replicate; Artificial intelligence; Machine learning; Software bug; Debugging; Layer (electronics); Process (computing); Code (set theory); Set (abstract data type); Deep learning; Data mining; Software; Source code; Programming language; Statistics; Mathematics","score_opus":0.0325642171348209,"score_gpt":0.32934321036649916,"score_spread":0.2967789932316783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2887758210","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6332984,0.0037152083,0.1846751,0.011253751,0.014986625,0.06951496,0.028722119,0.0052984245,0.048535462],"genre_scores_gemma":[0.74306923,0.000589839,0.11377128,0.0095997695,0.0015436077,0.09579622,0.013378894,0.001363458,0.020887671],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9531411,0.021873223,0.004529313,0.009779351,0.0094361855,0.0012408152],"domain_scores_gemma":[0.7946602,0.05955658,0.00877683,0.08693901,0.047007207,0.003060204],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.055847343,0.0016508249,0.0017438916,0.0013602407,0.002508232,0.0031834263,0.0027211145,0.00282681,0.012742512],"category_scores_gemma":[0.17112489,0.0008880651,0.0037826155,0.0016979374,0.0019843536,0.004215358,0.0030531005,0.0047102566,0.006277941],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.04520371,0.023291627,0.13434404,0.008437704,0.006051427,0.0046007275,0.019215126,0.01045563,0.0513814,0.035304453,0.17289038,0.48882383],"study_design_scores_gemma":[0.017314196,0.048656728,0.18541957,0.0026221022,0.0048622643,0.0032058016,0.008637465,0.021654628,0.044702243,0.054070204,0.6076962,0.0011586193],"about_ca_topic_score_codex":0.006041171,"about_ca_topic_score_gemma":0.0051597185,"teacher_disagreement_score":0.94415265,"about_ca_system_score_codex":0.0022121682,"about_ca_system_score_gemma":0.0052863136,"threshold_uncertainty_score":0.29535252},"labels":[],"label_agreement":null},{"id":"W2888049099","doi":"10.1007/s10664-018-9643-4","title":"Preventing duplicate bug reports by continuously querying bug reports","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Mitacs; Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Data deduplication; Software bug; BitTorrent tracker; Search engine indexing; Security bug; Information retrieval; Software; Database; Artificial intelligence; Cloud computing; Programming language","score_opus":0.014794330250701704,"score_gpt":0.2740173627620474,"score_spread":0.2592230325113457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2888049099","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5870996,0.006114238,0.34987333,0.0026452632,0.0007105131,0.0009826222,0.0037443158,0.042711478,0.0061185197],"genre_scores_gemma":[0.7615736,0.000798263,0.22943056,0.00045228578,0.00027922523,0.00015996677,0.0044342997,0.0008494577,0.002022242],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97736037,0.004419712,0.0022367102,0.004793025,0.0103961695,0.0007939152],"domain_scores_gemma":[0.85971427,0.06763268,0.026074817,0.023316788,0.020497924,0.0027635216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009665112,0.002384076,0.0025754205,0.010113718,0.0010165529,0.003865131,0.0047631366,0.0033334896,0.0013642854],"category_scores_gemma":[0.10799349,0.0012404303,0.0010119183,0.005081452,0.0007849063,0.005985355,0.0035241086,0.002490349,0.0014929232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00128154,0.0015134009,0.20277809,0.0011366296,0.0005844301,0.0006830449,0.0014251282,0.014181882,0.048715666,0.0021386854,0.018297106,0.7072645],"study_design_scores_gemma":[0.0005775551,0.0027840785,0.13184156,0.0004928637,0.0019281843,0.0047625476,0.0021578725,0.73313504,0.08256289,0.01753617,0.021802014,0.00041916926],"about_ca_topic_score_codex":0.0056443163,"about_ca_topic_score_gemma":0.006150316,"teacher_disagreement_score":0.010113718,"about_ca_system_score_codex":0.0007718369,"about_ca_system_score_gemma":0.003785508,"threshold_uncertainty_score":0.05111462},"labels":[],"label_agreement":null},{"id":"W2888214946","doi":"10.1145/3238147.3238194","title":"Generating reusable web components from mockups","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Mockup; Reusability; Web application; Web application development; Aspect-oriented programming; Software engineering; Programming language; Web service; Web modeling; World Wide Web; Human–computer interaction; Software; Engineering","score_opus":0.03089443352117113,"score_gpt":0.26673047394400723,"score_spread":0.2358360404228361,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2888214946","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10042877,0.00024191456,0.8766782,0.00011477201,0.0000742843,0.0010168208,0.00046418348,0.016663672,0.004317293],"genre_scores_gemma":[0.17581394,0.00017441408,0.816642,0.000071504,0.00001533437,0.0004807683,0.002120703,0.0016297393,0.0030515895],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99803454,0.00067801966,0.00015131547,0.00038867572,0.00066346105,0.00008398331],"domain_scores_gemma":[0.9880812,0.005118218,0.00075181463,0.0037003495,0.0020984788,0.00024981136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020054528,0.0017875037,0.00060169917,0.002898098,0.0006160693,0.0016226602,0.001320296,0.0012563075,0.002642907],"category_scores_gemma":[0.018169012,0.0009701008,0.0012596924,0.00081985747,0.00061631156,0.0012493107,0.0019571222,0.0008356001,0.001558803],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034922542,0.00045520594,0.00917107,0.0012262374,0.00013646614,0.0017957317,0.0028822282,0.04583838,0.11440843,0.010643667,0.010891228,0.802202],"study_design_scores_gemma":[0.00014568317,0.0006891571,0.013668418,0.0006377498,0.00024934614,0.0032541095,0.0014317487,0.5233179,0.35033605,0.024604548,0.081377886,0.00028749145],"about_ca_topic_score_codex":0.0010572709,"about_ca_topic_score_gemma":0.0018588444,"teacher_disagreement_score":0.002898098,"about_ca_system_score_codex":0.0005440844,"about_ca_system_score_gemma":0.0008592735,"threshold_uncertainty_score":0.010605991},"labels":[],"label_agreement":null},{"id":"W2888961423","doi":"","title":"Poster: Fast, Scalable and User-Guided Clone Detection","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Scalability; Source code; Transformation (genetics); Search engine indexing; Jaccard index; Data mining; Artificial intelligence; Database; Programming language; Pattern recognition (psychology); Biology","score_opus":0.0295438719690846,"score_gpt":0.28129037931804535,"score_spread":0.25174650734896076,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2888961423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037296206,0.0009137244,0.8563816,0.0018619082,0.0013057505,0.0005584692,0.004554992,0.09009837,0.0070289704],"genre_scores_gemma":[0.15183668,0.00048247617,0.7994338,0.00066292955,0.00062518986,0.00038200765,0.015685538,0.0059143687,0.024977023],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9967368,0.00047177996,0.00022582494,0.00073892536,0.0015971309,0.00022942187],"domain_scores_gemma":[0.9847317,0.0037016429,0.0006932153,0.004041677,0.005950798,0.00088107894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003286183,0.0012523349,0.0010338101,0.0027905072,0.0010860644,0.0029944626,0.001531181,0.0016004819,0.011163963],"category_scores_gemma":[0.0147304535,0.00062657957,0.0011493663,0.0022948135,0.00050309167,0.0032254232,0.0024798657,0.0017187013,0.010485366],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011713977,0.0003078272,0.010072972,0.0003685773,0.00016735827,0.00047553377,0.00040238927,0.004276995,0.09191373,0.004188235,0.16848639,0.7181686],"study_design_scores_gemma":[0.00042162766,0.00095958996,0.024939243,0.00013142747,0.00027599747,0.0033677644,0.0004088368,0.43976536,0.30822974,0.028349046,0.19276834,0.0003829993],"about_ca_topic_score_codex":0.0016858798,"about_ca_topic_score_gemma":0.002309818,"teacher_disagreement_score":0.011163963,"about_ca_system_score_codex":0.000657419,"about_ca_system_score_gemma":0.0011199771,"threshold_uncertainty_score":0.037347138},"labels":[],"label_agreement":null},{"id":"W2888962833","doi":"","title":"Poster: Designing Bug Detection Rules for Fewer False Alarms","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"False positive paradox; Computer science; Software bug; Constant false alarm rate; Open source; Static analysis; False alarm; True positive rate; False positives and false negatives; False positive rate; Data mining; Machine learning; Artificial intelligence; Software; Programming language","score_opus":0.038374997046337685,"score_gpt":0.2935458189583928,"score_spread":0.25517082191205515,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2888962833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09025367,0.00073725235,0.8714585,0.0028319228,0.0003780464,0.00087872596,0.0011097786,0.028921872,0.0034302117],"genre_scores_gemma":[0.32621565,0.00022386515,0.66555977,0.0010436468,0.00017499493,0.00039185624,0.0026642033,0.0016527323,0.0020732943],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9851521,0.0036101455,0.0020433243,0.0034912333,0.0049623647,0.00074070005],"domain_scores_gemma":[0.91882914,0.037458524,0.008203489,0.013894792,0.020076307,0.0015377888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014669601,0.0018998808,0.0013114385,0.0049737245,0.0012073393,0.0031125217,0.003267763,0.002507379,0.0024879447],"category_scores_gemma":[0.08199285,0.0011157953,0.0019536596,0.0015772202,0.0016740718,0.0043927496,0.0021455113,0.0034139906,0.0019119049],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009373903,0.0009071291,0.09283316,0.0008839223,0.0004302057,0.00061169465,0.00092103216,0.053921625,0.030230738,0.0075119603,0.034614723,0.7761963],"study_design_scores_gemma":[0.00030185838,0.0010039421,0.024051506,0.00042468857,0.00052269554,0.0017348367,0.0003407519,0.8402403,0.075558595,0.025665076,0.029856637,0.00029908714],"about_ca_topic_score_codex":0.003994431,"about_ca_topic_score_gemma":0.005983684,"teacher_disagreement_score":0.014669601,"about_ca_system_score_codex":0.0011683356,"about_ca_system_score_gemma":0.0030173324,"threshold_uncertainty_score":0.07758117},"labels":[],"label_agreement":null},{"id":"W2889872315","doi":"10.1145/3233027.3236393","title":"Teaching software product lines","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Snapshot (computer storage); Software engineering; Software; Product (mathematics); Engineering management; Engineering ethics; Engineering; Programming language; Operating system","score_opus":0.031069508928554604,"score_gpt":0.30437889012713265,"score_spread":0.27330938119857806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2889872315","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24161775,0.010770756,0.31966525,0.05591921,0.0053079682,0.00044758123,0.000954518,0.008715456,0.35660154],"genre_scores_gemma":[0.6313451,0.0108900415,0.15032719,0.0062121586,0.0010648157,0.00027856644,0.0014469462,0.0013171247,0.19711806],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99814594,0.0006199753,0.00008199717,0.00027744123,0.000685609,0.00018911877],"domain_scores_gemma":[0.9937518,0.0019410141,0.00035324175,0.00067490933,0.0016294098,0.001649588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019679468,0.00041819844,0.00019679229,0.000830931,0.00085550686,0.002842257,0.0008898859,0.0007983061,0.03151075],"category_scores_gemma":[0.009690898,0.00020350884,0.00026198092,0.0009723419,0.00081069587,0.0035023533,0.0026392797,0.0017011483,0.009817593],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000042633153,0.0007856913,0.0037391356,0.00038452185,0.000007639716,0.0002694511,0.0057124705,0.0007592542,0.0069022416,0.027478194,0.10865109,0.84526765],"study_design_scores_gemma":[0.000026739597,0.0004063341,0.0036666826,0.00051599887,0.00001110424,0.0006907867,0.006732096,0.002029312,0.011761419,0.046534583,0.92759645,0.000028599332],"about_ca_topic_score_codex":0.0007174884,"about_ca_topic_score_gemma":0.0018745468,"teacher_disagreement_score":0.03151075,"about_ca_system_score_codex":0.001482435,"about_ca_system_score_gemma":0.002545053,"threshold_uncertainty_score":0.105413914},"labels":[],"label_agreement":null},{"id":"W2889927216","doi":"","title":"[Journal First] An Empirical Study of Early Access Games on the Steam Platform","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Game Developer; Empirical research; Popularity; World Wide Web; Multimedia; Game design","score_opus":0.10599219869073978,"score_gpt":0.36780825437724596,"score_spread":0.2618160556865062,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2889927216","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96303356,0.00026016735,0.0006125247,0.0009324133,0.00006393796,0.00021192261,0.0006846134,0.000014673382,0.034186125],"genre_scores_gemma":[0.9906557,0.00032545414,0.00065449707,0.0006361366,0.00004069855,0.0002528314,0.0007122449,0.000020785357,0.0067016315],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99673444,0.0012412741,0.0003038385,0.00027719003,0.0010639053,0.0003793565],"domain_scores_gemma":[0.95182586,0.022129165,0.0103089,0.002267387,0.009219712,0.004249018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035608574,0.00024154418,0.00037507227,0.0020806727,0.001768116,0.0028488745,0.0011202578,0.0008234411,0.014645643],"category_scores_gemma":[0.047067415,0.00033838436,0.0002808304,0.0025721753,0.001545484,0.0037533166,0.0015490421,0.001960132,0.0023733035],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022928174,0.0022782194,0.88872063,0.0003976543,0.00007701532,0.00049287686,0.03597283,0.00019976446,0.0005443979,0.0064120493,0.01498325,0.04969217],"study_design_scores_gemma":[0.00004404146,0.0005082,0.93136704,0.0002661522,0.00003002144,0.00026194818,0.041075125,0.0006634685,0.0002989161,0.0008208407,0.0246216,0.000042694097],"about_ca_topic_score_codex":0.012495259,"about_ca_topic_score_gemma":0.0149205765,"teacher_disagreement_score":0.014645643,"about_ca_system_score_codex":0.0011532733,"about_ca_system_score_gemma":0.001564613,"threshold_uncertainty_score":0.04899454},"labels":[],"label_agreement":null},{"id":"W2890502052","doi":"","title":"[Journal First] Measuring Program Comprehension: A Large-Scale Field Study with Professionals","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Program comprehension; Comprehension; Computer science; Field (mathematics); Software engineering; World Wide Web; Software; Multimedia; Human–computer interaction; Programming language; Software system","score_opus":0.04064767116405217,"score_gpt":0.3180627877946147,"score_spread":0.27741511663056256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2890502052","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9817355,0.0002336702,0.008085468,0.00050120335,0.000078361736,0.0007860203,0.00085302454,0.00018121603,0.007545489],"genre_scores_gemma":[0.9803726,0.00027876798,0.011192313,0.00070703164,0.00010540233,0.0013151759,0.0014745244,0.000084916035,0.004469334],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9979285,0.0008248202,0.00019509254,0.00038026768,0.0005439763,0.00012726772],"domain_scores_gemma":[0.96513945,0.016165564,0.0034294762,0.0024534152,0.010641296,0.0021707406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004069919,0.00034155542,0.0003420858,0.0025087937,0.0013715018,0.001029928,0.00073883997,0.00069755234,0.0038742684],"category_scores_gemma":[0.025153292,0.0003901664,0.00026065944,0.0016974725,0.00074052927,0.0017127417,0.00093011616,0.0008534156,0.0012678215],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000656954,0.006570127,0.5424671,0.00094659545,0.00013384293,0.000782031,0.08791394,0.00047775218,0.01632673,0.00094572164,0.02720057,0.31557882],"study_design_scores_gemma":[0.00021586877,0.0052180006,0.91351366,0.00027457453,0.00010662743,0.0009984452,0.028965842,0.0017732894,0.0071675554,0.0012411118,0.040415153,0.00010989817],"about_ca_topic_score_codex":0.002560257,"about_ca_topic_score_gemma":0.004911869,"teacher_disagreement_score":0.004069919,"about_ca_system_score_codex":0.0005205952,"about_ca_system_score_gemma":0.0011026127,"threshold_uncertainty_score":0.021524072},"labels":[],"label_agreement":null},{"id":"W2890687291","doi":"","title":"[Journal First] Are Fix-Inducing Changes a Moving Target?: A Longitudinal Case Study of Just-in-Time Defect Prediction","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Code (set theory); Brier score; Interpretation (philosophy); Calibration; Artificial intelligence; Statistics; Mathematics; Programming language","score_opus":0.050354734497892885,"score_gpt":0.2965110632208014,"score_spread":0.24615632872290852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2890687291","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9870603,0.00030258874,0.0085660415,0.0012871126,0.00005300501,0.00007801202,0.0008532564,0.00020270472,0.0015969443],"genre_scores_gemma":[0.9933501,0.0001288447,0.0041241413,0.00016970331,0.000032962023,0.000039676986,0.0012500699,0.00005995442,0.0008444988],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9968995,0.0011314915,0.00023815414,0.0006845146,0.0008074596,0.00023887042],"domain_scores_gemma":[0.95492375,0.020368215,0.007125095,0.007740305,0.0078978855,0.0019447156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070668934,0.00036009194,0.00042390777,0.0016035319,0.0010790661,0.0014687818,0.0012336341,0.0013373265,0.0016502709],"category_scores_gemma":[0.050896667,0.00033659593,0.0005789544,0.0014176064,0.0010588609,0.002704854,0.0011792352,0.0015603602,0.0007315044],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000364771,0.0007151933,0.8796367,0.00009961771,0.00014774107,0.0014927597,0.0033861983,0.019950802,0.0013892994,0.002148672,0.007861974,0.082806304],"study_design_scores_gemma":[0.00010435469,0.0024171907,0.61548793,0.000272073,0.00018683563,0.0035208159,0.004865294,0.33356056,0.0070322934,0.013225763,0.01908208,0.00024478877],"about_ca_topic_score_codex":0.014422948,"about_ca_topic_score_gemma":0.017191019,"teacher_disagreement_score":0.014422948,"about_ca_system_score_codex":0.0013013509,"about_ca_system_score_gemma":0.00096659776,"threshold_uncertainty_score":0.03737372},"labels":[],"label_agreement":null},{"id":"W2890814105","doi":"","title":"[Journal First] Inference of Development Activities from Interaction with Uninstrumented Applications","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Generalizability theory; Classifier (UML); Inference; Machine learning; Software; Data mining; Support vector machine; Software engineering; Data science; Artificial intelligence","score_opus":0.027218466849824794,"score_gpt":0.28296712316630507,"score_spread":0.25574865631648025,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2890814105","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13054478,0.0014890293,0.82651085,0.0010459927,0.00035314178,0.0003407106,0.013656775,0.015399722,0.010659094],"genre_scores_gemma":[0.6696373,0.0004980026,0.30735543,0.0001890807,0.00018309228,0.00033408738,0.015808512,0.00060875725,0.0053857793],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9981536,0.0004950257,0.000111986876,0.00065570086,0.00047432774,0.00010941327],"domain_scores_gemma":[0.99104846,0.004216975,0.0010440055,0.0021650954,0.0012839142,0.00024141444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002004807,0.0006871651,0.0005616452,0.0029486655,0.0004911777,0.0009775329,0.0016286985,0.0011043798,0.00442928],"category_scores_gemma":[0.016589567,0.00050490233,0.0011378207,0.0023264429,0.0005495632,0.0016213801,0.0007860224,0.0011743467,0.0025633348],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006129416,0.00042818312,0.10269972,0.00064809254,0.00021733115,0.00039635735,0.00060697406,0.059866298,0.0071118274,0.00790103,0.038632948,0.7808783],"study_design_scores_gemma":[0.00006746064,0.00019200712,0.07695551,0.00013985104,0.000092723545,0.0005292136,0.00013104928,0.86822915,0.0075891674,0.017987156,0.027992746,0.00009389459],"about_ca_topic_score_codex":0.012017877,"about_ca_topic_score_gemma":0.013862064,"teacher_disagreement_score":0.012017877,"about_ca_system_score_codex":0.000630311,"about_ca_system_score_gemma":0.0013323576,"threshold_uncertainty_score":0.02389586},"labels":[],"label_agreement":null},{"id":"W2892077113","doi":"","title":"[Journal First] Analyzing the Effects of Test Driven Development in GitHub","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Agile software development; Test-driven development; Computer science; Commit; Software engineering; Software quality; Quality (philosophy); Software; Software development process; Software development; Operating system; Database","score_opus":0.017641440733456477,"score_gpt":0.2635392492206144,"score_spread":0.24589780848715792,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2892077113","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91869986,0.0029230693,0.01804778,0.002730134,0.00078300695,0.0013724741,0.016407015,0.0029920726,0.036044497],"genre_scores_gemma":[0.9515073,0.0009865441,0.0171204,0.0014402929,0.0003335652,0.0015275524,0.017131865,0.0010951126,0.008857386],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99177986,0.0023990544,0.0005840236,0.0007661157,0.0040692347,0.00040173516],"domain_scores_gemma":[0.88383967,0.049500395,0.017961975,0.018029993,0.027188174,0.0034797739],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005683177,0.00059872586,0.00060445064,0.0046962653,0.00083343836,0.0026550398,0.001624042,0.0007888599,0.007839479],"category_scores_gemma":[0.07356545,0.0003810606,0.0012292414,0.0060885837,0.0011451342,0.0017879639,0.0021802492,0.0011960338,0.0024781052],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044723395,0.0020232843,0.4387167,0.0031952783,0.001486596,0.0006850406,0.0036581678,0.008554527,0.009965767,0.005207965,0.07831637,0.44371787],"study_design_scores_gemma":[0.00031407652,0.0030616499,0.9216644,0.0005934967,0.00067143014,0.00052067003,0.0019890782,0.011918342,0.007266191,0.003225609,0.048615657,0.00015923806],"about_ca_topic_score_codex":0.01300644,"about_ca_topic_score_gemma":0.013349552,"teacher_disagreement_score":0.01300644,"about_ca_system_score_codex":0.002291641,"about_ca_system_score_gemma":0.0020741452,"threshold_uncertainty_score":0.03005588},"labels":[],"label_agreement":null},{"id":"W2892136838","doi":"","title":"[Journal First] On the Use of Hidden Markov Model to Predict the Time to Fix Bugs","year":2018,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Hidden Markov model; Computer science; Software bug; Software regression; Software; Markov model; Predictive modelling; Data science; Software engineering; Markov chain; Data mining; Machine learning; Software development; Artificial intelligence; Software quality; Programming language","score_opus":0.05873565329723863,"score_gpt":0.27639442820363846,"score_spread":0.21765877490639984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2892136838","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059314713,0.00939244,0.9025846,0.007089887,0.001819346,0.00023115549,0.0033345919,0.00453504,0.01169817],"genre_scores_gemma":[0.53619975,0.009600612,0.41904715,0.0029227412,0.0018171683,0.00032973476,0.0040222704,0.00090671144,0.025153833],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99854976,0.0004548685,0.00016128279,0.00035272737,0.0003784075,0.0001029255],"domain_scores_gemma":[0.98547494,0.010227098,0.00067654153,0.001178371,0.00214087,0.00030215087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002690313,0.0013207707,0.0012394163,0.004047393,0.0007371774,0.0017929007,0.0017279702,0.0026492493,0.003961299],"category_scores_gemma":[0.024869049,0.001050571,0.0021046475,0.003345771,0.0005253968,0.0025315785,0.0011585386,0.0025119581,0.002518229],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010549632,0.00079414994,0.09127492,0.0011062343,0.0013033941,0.0008857811,0.00131561,0.15958291,0.004032943,0.020648077,0.0462295,0.6717715],"study_design_scores_gemma":[0.00013623436,0.00059092575,0.017891211,0.00041789774,0.00067993026,0.00078045996,0.00028016165,0.9133453,0.0050984775,0.027771529,0.032715883,0.00029204585],"about_ca_topic_score_codex":0.032673586,"about_ca_topic_score_gemma":0.02896663,"teacher_disagreement_score":0.032673586,"about_ca_system_score_codex":0.00084851525,"about_ca_system_score_gemma":0.0019770223,"threshold_uncertainty_score":0.06496686},"labels":[],"label_agreement":null},{"id":"W2894722607","doi":"10.1145/3239235.3239525","title":"An empirical study of design discussions in code review","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Japan Society for the Promotion of Science","keywords":"Computer science; Code (set theory); Software engineering; Focus (optics); Empirical research; Software quality; Software design; Software; Code review; Programming language; Software development; Epistemology","score_opus":0.10240411612980647,"score_gpt":0.4161624724902682,"score_spread":0.31375835636046173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2894722607","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9782065,0.0019132496,0.004123826,0.0029142082,0.00012207666,0.0013755398,0.00018696772,0.000059451573,0.0110981865],"genre_scores_gemma":[0.9914483,0.000863589,0.0033034014,0.0011894462,0.00008124812,0.0016572648,0.00018976034,0.00006001497,0.001207105],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.72413445,0.21922904,0.013454069,0.0082212025,0.030670596,0.004290706],"domain_scores_gemma":[0.07422322,0.7969522,0.08051897,0.010546272,0.031871654,0.0058876565],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.1494695,0.0006870993,0.0010002971,0.007879945,0.008473276,0.010519984,0.0035639398,0.0040169037,0.0040961676],"category_scores_gemma":[0.65712357,0.0014259361,0.0007696266,0.006590031,0.010370237,0.012237874,0.008149331,0.006964841,0.0008612851],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010267566,0.0018567682,0.16730684,0.0030689035,0.00014894757,0.001039492,0.7574213,0.00029238767,0.0010225634,0.0050547747,0.0041543394,0.05760695],"study_design_scores_gemma":[0.0006288085,0.0030861984,0.24037763,0.0055323965,0.00023289761,0.0016923232,0.68372697,0.0030588247,0.0018358367,0.006287228,0.05318095,0.00035987786],"about_ca_topic_score_codex":0.0042271377,"about_ca_topic_score_gemma":0.004957409,"teacher_disagreement_score":0.8505305,"about_ca_system_score_codex":0.0109882075,"about_ca_system_score_gemma":0.014919603,"threshold_uncertainty_score":0.79047966},"labels":[],"label_agreement":null},{"id":"W2895482093","doi":"10.1002/spe.2639","title":"A systematic literature review on the detection of smells and their evolution in object‐oriented and service‐oriented systems","year":2018,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Concordia University","funders":"","keywords":"Code smell; Computer science; Software engineering; Source code; Identification (biology); Suite; Service (business); Data science; Information retrieval; Software; Software development; Programming language; Software quality","score_opus":0.010999122528343468,"score_gpt":0.2692355839895828,"score_spread":0.25823646146123935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2895482093","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019102466,0.9956535,0.0004945248,0.00049061555,0.00013951493,0.00029096074,0.0004183438,0.000015611622,0.00058669],"genre_scores_gemma":[0.013944297,0.9822204,0.002167456,0.0004931736,0.00007111276,0.00053389143,0.0003867217,0.000011486313,0.0001714083],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.98754996,0.003802289,0.0049388036,0.0008685314,0.0025160254,0.00032441155],"domain_scores_gemma":[0.93960327,0.04218437,0.008446821,0.001031864,0.008229154,0.00050451735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011849892,0.001587537,0.0044027558,0.028357241,0.0009526205,0.0031454964,0.0017166917,0.0017503627,0.003489849],"category_scores_gemma":[0.05951915,0.0011166597,0.005490303,0.022053018,0.0011691923,0.0042982125,0.0021782937,0.0012795139,0.0004713327],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010435878,0.000030102343,0.00151305,0.87028146,0.0019205175,0.00022435017,0.0010448586,0.00017704198,0.00041180532,0.0008232498,0.0033835438,0.12008568],"study_design_scores_gemma":[0.000043163403,0.00016224473,0.0047529787,0.94056195,0.0099706445,0.00048795866,0.0011322134,0.000097199205,0.0003233664,0.000561664,0.041860603,0.000045952856],"about_ca_topic_score_codex":0.00873784,"about_ca_topic_score_gemma":0.025843568,"teacher_disagreement_score":0.028357241,"about_ca_system_score_codex":0.003915984,"about_ca_system_score_gemma":0.018516947,"threshold_uncertainty_score":0.06266898},"labels":[],"label_agreement":null},{"id":"W2895720364","doi":"10.1145/3239372.3239406","title":"Recommending Model Refactoring Rules from Refactoring Examples","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Code refactoring; Computer science; Programming language; Software engineering; Class (philosophy); Code (set theory); Metamodeling; Artificial intelligence; Software","score_opus":0.08260532393748822,"score_gpt":0.3091265865348926,"score_spread":0.2265212625974044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2895720364","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2158276,0.0021918593,0.76262814,0.0010506861,0.00016155277,0.00067816506,0.0033106885,0.010121691,0.0040297224],"genre_scores_gemma":[0.34591576,0.0008955935,0.6404226,0.00029160763,0.0000781155,0.00033747102,0.009611923,0.0005377411,0.0019091871],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975567,0.00067200186,0.00024022513,0.00072687963,0.00070301717,0.000101197744],"domain_scores_gemma":[0.9822099,0.009695196,0.0011326538,0.0022134385,0.0043751462,0.00037370712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029230856,0.0018098138,0.0011707948,0.0045695677,0.0005307948,0.0016142981,0.0020418484,0.0018224309,0.0013336169],"category_scores_gemma":[0.023664955,0.0006715997,0.0011547337,0.0016529929,0.0003798862,0.0017159946,0.0007334715,0.0015835862,0.0012822396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043965044,0.00094620156,0.04690675,0.0012924713,0.00038176845,0.000681725,0.0008379973,0.09514034,0.012807847,0.0028110794,0.030775938,0.8069782],"study_design_scores_gemma":[0.00015388886,0.00019203483,0.0068266904,0.00036968797,0.0003798408,0.00042509474,0.00038456297,0.95255184,0.017436896,0.0060700127,0.015127797,0.000081709026],"about_ca_topic_score_codex":0.008460223,"about_ca_topic_score_gemma":0.022378568,"teacher_disagreement_score":0.008460223,"about_ca_system_score_codex":0.0006965293,"about_ca_system_score_gemma":0.0022086967,"threshold_uncertainty_score":0.01682198},"labels":[],"label_agreement":null},{"id":"W2896266317","doi":"10.1109/aire.2018.00008","title":"User Feedback from Tweets vs App Store Reviews: An Exploratory Study of Frequency, Timing and Content","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Latent Dirichlet allocation; Social media; Context (archaeology); Feature (linguistics); Set (abstract data type); App store; Sentiment analysis; Mobile apps; World Wide Web; Topic model; Exploratory research; User-generated content; Information retrieval; Natural language processing","score_opus":0.11039948118147609,"score_gpt":0.3110871775228645,"score_spread":0.20068769634138844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2896266317","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9973132,0.00009125678,0.0006859421,0.00006401579,0.000009313429,0.00014397826,0.0008747468,0.000033146465,0.0007843962],"genre_scores_gemma":[0.9962483,0.0001079692,0.0016164081,0.00008773348,0.000027069224,0.00039255255,0.00068816415,0.000023281425,0.00080856564],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99800605,0.0008717734,0.00014797982,0.0002399234,0.000577273,0.000156985],"domain_scores_gemma":[0.97307736,0.01924403,0.0034322145,0.00062939717,0.0030462854,0.00057071703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020687554,0.00030162864,0.00043688543,0.0017403373,0.000579549,0.0011484076,0.00027020104,0.00046533393,0.0012038946],"category_scores_gemma":[0.017790603,0.00020830015,0.00027683933,0.0015364834,0.00035872447,0.001451024,0.0007875301,0.00058191916,0.0006722211],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020791749,0.001199969,0.7621546,0.0016852218,0.00022045216,0.0013816025,0.09880041,0.0005525263,0.022077074,0.0005223796,0.004810526,0.10451604],"study_design_scores_gemma":[0.000035183846,0.0020618811,0.93021053,0.00021899893,0.00015082858,0.0008382287,0.045101043,0.004593362,0.008132593,0.0003272199,0.008226323,0.0001037594],"about_ca_topic_score_codex":0.0015467993,"about_ca_topic_score_gemma":0.002336167,"teacher_disagreement_score":0.0020687554,"about_ca_system_score_codex":0.00038645408,"about_ca_system_score_gemma":0.0003248676,"threshold_uncertainty_score":0.01094079},"labels":[],"label_agreement":null},{"id":"W2896375752","doi":"10.1145/3273934.3273937","title":"An Empirical Study of Metric-based Comparisons of Software Libraries","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Canada Excellence Research Chairs, Government of Canada","keywords":"Computer science; Metric (unit); Software metric; Popularity; Java; Domain (mathematical analysis); Set (abstract data type); Software; World Wide Web; Empirical research; Software engineering; Software development; Information retrieval; Data science; Software quality; Programming language","score_opus":0.0559064641556428,"score_gpt":0.35021473765816374,"score_spread":0.2943082735025209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2896375752","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9807829,0.00085199164,0.012976271,0.00034490708,0.000033076,0.0003947646,0.0010422058,0.00016949832,0.0034043724],"genre_scores_gemma":[0.9840894,0.00017845267,0.0137673495,0.00007553178,0.000022514245,0.0005253709,0.0010604988,0.000045195087,0.00023567729],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9209955,0.04668026,0.007496148,0.0045932457,0.019116195,0.0011187154],"domain_scores_gemma":[0.34557462,0.5199624,0.06621226,0.0148540605,0.05003687,0.003359792],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.053723477,0.00061107497,0.0006329951,0.010368299,0.00089511054,0.002485113,0.0013958527,0.00085189834,0.0011533684],"category_scores_gemma":[0.34236473,0.00037095536,0.0006825829,0.014011675,0.0014462961,0.006471674,0.0021838008,0.0013900078,0.00033288248],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007275831,0.0010699313,0.8458134,0.0018678368,0.0005217729,0.00016910626,0.011059942,0.0019159664,0.0016217792,0.002870593,0.0030898948,0.1292723],"study_design_scores_gemma":[0.00015384862,0.0030347076,0.9298781,0.00074498606,0.00028273897,0.00051470875,0.016226023,0.028648993,0.004478224,0.0034376003,0.012428662,0.00017145173],"about_ca_topic_score_codex":0.0017612803,"about_ca_topic_score_gemma":0.0027219132,"teacher_disagreement_score":0.9896317,"about_ca_system_score_codex":0.0019259123,"about_ca_system_score_gemma":0.0015654713,"threshold_uncertainty_score":0.28412026},"labels":[],"label_agreement":null},{"id":"W2896543198","doi":"10.1002/smr.2117","title":"Evaluating filter fuzzy analogy homogenous ensembles for software development effort estimation","year":2018,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Analogy; Fuzzy logic; Defuzzification; Artificial intelligence; Computer science; Filter (signal processing); Fuzzy set; Machine learning; Neuro-fuzzy; Feature (linguistics); Fuzzy classification; Data mining; Mathematics; Fuzzy number; Fuzzy control system","score_opus":0.04033265670826348,"score_gpt":0.3405747214067008,"score_spread":0.30024206469843734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2896543198","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64702964,0.00052746775,0.3497383,0.00011042446,0.000034838842,0.000095949414,0.00021227462,0.00036787274,0.0018831824],"genre_scores_gemma":[0.9470901,0.00008157296,0.052229878,0.000018741519,0.000013964474,0.000054577795,0.0002243585,0.0000139307285,0.00027288587],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977477,0.0009332779,0.0001341908,0.00043534767,0.0006359877,0.00011347916],"domain_scores_gemma":[0.9857451,0.010514026,0.0008590513,0.001182838,0.0015133244,0.00018574603],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048419964,0.0007349832,0.00089529273,0.0032799826,0.00037876153,0.00092851365,0.000650588,0.0008415382,0.0008163513],"category_scores_gemma":[0.023558049,0.00022762144,0.0007966627,0.0017964457,0.00024587158,0.0017678603,0.000803629,0.000636137,0.0001865903],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041416366,0.00028576446,0.04609771,0.00015127612,0.0006289528,0.000078580175,0.00029991462,0.6058815,0.0053322143,0.0030941558,0.00093622773,0.33679947],"study_design_scores_gemma":[0.000007744522,0.00017318559,0.010623326,0.000018442835,0.000066130466,0.000030984265,0.000068723435,0.9833165,0.0028766273,0.0023964876,0.0004027984,0.000019101655],"about_ca_topic_score_codex":0.0029949339,"about_ca_topic_score_gemma":0.0030315781,"teacher_disagreement_score":0.0048419964,"about_ca_system_score_codex":0.00075603684,"about_ca_system_score_gemma":0.000571274,"threshold_uncertainty_score":0.025607228},"labels":[],"label_agreement":null},{"id":"W2896875458","doi":"10.1109/re.2018.00033","title":"Morse: Reducing the Feature Interaction Explosion Problem using Subject Matter Knowledge as Abstract Requirements","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Critical Systems Labs","funders":"","keywords":"Feature (linguistics); Computer science; Domain (mathematical analysis); Risk analysis (engineering); Simple (philosophy); Subject-matter expert; Point (geometry); Artificial intelligence; Human–computer interaction; Data science; Expert system; Mathematics","score_opus":0.05720230423770137,"score_gpt":0.3512630367988539,"score_spread":0.2940607325611525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2896875458","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016309276,0.00006722781,0.976977,0.00031442236,0.000027663536,0.00026933808,0.0002634461,0.0029752427,0.002796308],"genre_scores_gemma":[0.10844347,0.000105740386,0.88689286,0.00026167202,0.000034774435,0.00028234595,0.0008921229,0.0007367202,0.002350322],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9900073,0.0026532386,0.00058994174,0.001315937,0.0048238556,0.00060964574],"domain_scores_gemma":[0.9676982,0.021299303,0.0019853518,0.0060356734,0.0025922183,0.0003893292],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006000444,0.0017415474,0.0014719645,0.0033135249,0.0010691365,0.002270971,0.0025956884,0.0018623493,0.0059648617],"category_scores_gemma":[0.04032008,0.0015635927,0.0043534525,0.0012830897,0.002538853,0.00587549,0.0059716143,0.003980067,0.0011290363],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070684624,0.00082042493,0.008377445,0.0013733911,0.00043029746,0.0023091114,0.0019613965,0.18501815,0.030707238,0.17541945,0.010740676,0.5821356],"study_design_scores_gemma":[0.00015811292,0.00032508266,0.0015965704,0.00024299246,0.00025491146,0.0006583586,0.0005385944,0.69152546,0.03580757,0.24757375,0.021192053,0.00012650651],"about_ca_topic_score_codex":0.0037065835,"about_ca_topic_score_gemma":0.005698486,"teacher_disagreement_score":0.006000444,"about_ca_system_score_codex":0.00122914,"about_ca_system_score_gemma":0.002764288,"threshold_uncertainty_score":0.03173375},"labels":[],"label_agreement":null},{"id":"W2897662303","doi":"10.1007/s10664-018-9656-z","title":"High-level software requirements and iteration changes: a predictive model","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Systems, Applications & Products in Data Processing (Canada); University of Victoria","funders":"","keywords":"Computer science; Software engineering; Reliability engineering; Engineering","score_opus":0.05761095084229382,"score_gpt":0.3005239212894088,"score_spread":0.242912970447115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2897662303","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90352726,0.00042181058,0.08411842,0.0012217442,0.00006325852,0.00013591466,0.0006967203,0.00069194974,0.009123077],"genre_scores_gemma":[0.9949923,0.00011828767,0.0030490926,0.000039847782,0.000016824055,0.000045920802,0.00027955853,0.000040037405,0.0014181362],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985355,0.00058683095,0.00006205454,0.00032132625,0.00030904062,0.0001852555],"domain_scores_gemma":[0.95819336,0.03492821,0.0019644476,0.002219778,0.0020431874,0.0006510645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00424396,0.0012235928,0.00092181476,0.002590059,0.00068312534,0.003018691,0.0029327006,0.002491886,0.0057912865],"category_scores_gemma":[0.034580328,0.0010906435,0.0015381845,0.0020302914,0.0014615301,0.0041576945,0.0012201297,0.002728245,0.0016983738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001097263,0.0018962864,0.18506393,0.00019719142,0.00038595725,0.000593857,0.0013549557,0.7164422,0.0011167426,0.022722725,0.0034497255,0.06567913],"study_design_scores_gemma":[0.000032055148,0.00009868672,0.01517061,0.00002330983,0.00008614695,0.000090132155,0.00010276857,0.97434366,0.00020056892,0.009611786,0.00021905027,0.00002129125],"about_ca_topic_score_codex":0.017453218,"about_ca_topic_score_gemma":0.011113468,"teacher_disagreement_score":0.017453218,"about_ca_system_score_codex":0.0016054567,"about_ca_system_score_gemma":0.0016623931,"threshold_uncertainty_score":0.034703255},"labels":[],"label_agreement":null},{"id":"W2897717013","doi":"10.1002/smr.2114","title":"Support vector regression‐based imputation in analogy‐based software development effort estimation","year":2018,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Imputation (statistics); Analogy; Missing data; Computer science; Support vector machine; Data mining; Artificial intelligence; Regression; Euclidean distance; Machine learning; Statistics; Mathematics","score_opus":0.015284866129832408,"score_gpt":0.29812949750248796,"score_spread":0.28284463137265553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2897717013","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18296948,0.00068271265,0.81215805,0.00040007214,0.00006912271,0.00013012426,0.0006138835,0.001558243,0.0014183053],"genre_scores_gemma":[0.7945087,0.00020467215,0.2031123,0.00008985086,0.00003286161,0.00014736207,0.0010350518,0.000066122724,0.0008030973],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9931589,0.0039533535,0.0004938812,0.0010377963,0.0011570993,0.00019894273],"domain_scores_gemma":[0.97371066,0.01791263,0.002660842,0.0022751922,0.0031842142,0.0002564913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008232255,0.00083912036,0.0012906147,0.00241706,0.00039543872,0.0010186499,0.0016213808,0.00090514764,0.0014892869],"category_scores_gemma":[0.040151574,0.00038806262,0.0011153893,0.0023665314,0.0004276178,0.0020164405,0.0011558001,0.0019151544,0.0005670411],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004029267,0.00045072127,0.04888167,0.00043337402,0.00033251048,0.00019259762,0.00041581632,0.6037804,0.001999437,0.006815239,0.0028927042,0.3334026],"study_design_scores_gemma":[0.0000130958015,0.000110334244,0.0035803793,0.000040941617,0.000019299983,0.000030453237,0.000061819585,0.98963577,0.0013423996,0.004448373,0.00069822464,0.000018915971],"about_ca_topic_score_codex":0.0029993653,"about_ca_topic_score_gemma":0.002619296,"teacher_disagreement_score":0.008232255,"about_ca_system_score_codex":0.0005683112,"about_ca_system_score_gemma":0.0009544571,"threshold_uncertainty_score":0.043536842},"labels":[],"label_agreement":null},{"id":"W2897892445","doi":"10.1109/re.2018.00-23","title":"Crowd-Focused Semi-Automated Requirements Engineering for Evolution Towards Sustainability","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Engineering and Physical Sciences Research Council; Natural Sciences and Engineering Research Council of Canada; Agència per a la Competitivitat de l’Empresa; Generalitat de Catalunya; European Commission","keywords":"Sustainability; Computer science; Requirements engineering; Requirements elicitation; Work (physics); Negotiation; Product-service system; Social sustainability; Product (mathematics); Systems engineering; Process management; Knowledge management; Engineering; Software","score_opus":0.020923371338112673,"score_gpt":0.30510104139037747,"score_spread":0.2841776700522648,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2897892445","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01927941,0.00007864085,0.9731794,0.0007285829,0.000041090607,0.0010099886,0.00013505046,0.001141007,0.0044068214],"genre_scores_gemma":[0.17342788,0.000108337044,0.8220937,0.00025019373,0.000026533797,0.0011882543,0.00048436332,0.000274113,0.0021466895],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9800806,0.013723728,0.0009541152,0.0012736012,0.003442802,0.00052517414],"domain_scores_gemma":[0.9551906,0.02935053,0.002483519,0.0070553483,0.004621861,0.0012981282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014784011,0.0012475216,0.00079252664,0.0015362948,0.0017696341,0.0032210986,0.0025901343,0.0023062145,0.005407523],"category_scores_gemma":[0.03633663,0.0011140656,0.00197715,0.0007382871,0.0031265465,0.0037139656,0.007377937,0.0022872277,0.0014002003],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058574905,0.0009481812,0.004076544,0.0019559474,0.0002426529,0.00352117,0.023982672,0.42133474,0.08365588,0.23468846,0.011956239,0.21305183],"study_design_scores_gemma":[0.00018313715,0.00033833104,0.0012202607,0.00043658036,0.00006574722,0.0007428678,0.00407983,0.7213309,0.021828994,0.17652094,0.07308546,0.00016698091],"about_ca_topic_score_codex":0.00286427,"about_ca_topic_score_gemma":0.0040743826,"teacher_disagreement_score":0.014784011,"about_ca_system_score_codex":0.0021831193,"about_ca_system_score_gemma":0.0046959175,"threshold_uncertainty_score":0.078186214},"labels":[],"label_agreement":null},{"id":"W2898370077","doi":"10.1145/3278122.3278127","title":"Exploring feature interactions without specifications: a controlled experiment","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Feature (linguistics); Graph; Control flow; Data mining; Control flow graph; Information flow; Machine learning; Artificial intelligence; Theoretical computer science; Programming language","score_opus":0.17213516854012698,"score_gpt":0.3379760282863929,"score_spread":0.16584085974626592,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898370077","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98202896,0.00016458878,0.01112289,0.00012762043,0.00011147409,0.004428819,0.00052477786,0.00039681655,0.0010940742],"genre_scores_gemma":[0.9241345,0.0002636073,0.05415518,0.00071390154,0.00017191462,0.0149141615,0.0010732425,0.00021197669,0.0043614735],"study_design_codex":"bench_or_experimental","study_design_gemma":"randomized_trial","domain_scores_codex":[0.99636203,0.0016832001,0.0003137146,0.0009077261,0.00041273914,0.00032057305],"domain_scores_gemma":[0.9361268,0.053966492,0.0028736575,0.0031754388,0.0021515395,0.0017059891],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006336342,0.001864579,0.0010628348,0.00062372675,0.00086731213,0.001145001,0.002030787,0.00213553,0.00910437],"category_scores_gemma":[0.022262868,0.0009372275,0.00073779974,0.00038188248,0.0013175362,0.001915374,0.0014534212,0.0015636159,0.0014788694],"study_design_candidate":"randomized_trial","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.09270469,0.26224396,0.035158053,0.009141103,0.00088791037,0.0029193268,0.031218432,0.012059704,0.34568825,0.0031152612,0.011094672,0.19376859],"study_design_scores_gemma":[0.064483814,0.5468582,0.10452758,0.0009140912,0.002241747,0.002188339,0.009802209,0.054949414,0.1656755,0.012604496,0.034875024,0.0008796424],"about_ca_topic_score_codex":0.00055572804,"about_ca_topic_score_gemma":0.0007118864,"teacher_disagreement_score":0.00910437,"about_ca_system_score_codex":0.00034666184,"about_ca_system_score_gemma":0.0008217056,"threshold_uncertainty_score":0.033510208},"labels":[],"label_agreement":null},{"id":"W2898501178","doi":"10.1109/vlhcc.2018.8506491","title":"ZenStates: Easy-to-Understand Yet Expressive Specifications for Creative Interactive Environments","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Blackboard (design pattern); Human–computer interaction; Interface (matter); Expressive power; User interface; Software engineering; Software; Multimedia; Programming language","score_opus":0.07928255459081862,"score_gpt":0.3204786716394142,"score_spread":0.24119611704859556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898501178","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0040432676,0.000056836714,0.98934513,0.00018523948,0.000022185035,0.000097647026,0.00032070797,0.0037977607,0.0021311806],"genre_scores_gemma":[0.13540144,0.0002452565,0.85549855,0.00023541253,0.000023482295,0.000669396,0.0012324927,0.001608561,0.005085389],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9977374,0.00095304725,0.00025999683,0.00023925937,0.0006703536,0.00013996848],"domain_scores_gemma":[0.9949457,0.0031892727,0.00037452485,0.00097398943,0.00037058326,0.00014593208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002478031,0.0011469371,0.00047735422,0.0006455299,0.0006363646,0.0032116629,0.002113679,0.0013982899,0.007884241],"category_scores_gemma":[0.009260407,0.001132309,0.0014479148,0.00045753524,0.0020804477,0.0037905534,0.0025123965,0.0022420145,0.0018711452],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004398294,0.00015023729,0.0021315638,0.0007981702,0.00007678091,0.0003560589,0.0028850113,0.10881535,0.01530591,0.75422984,0.012528375,0.10228277],"study_design_scores_gemma":[0.00016091096,0.0001886223,0.0004352861,0.0003916027,0.000094745796,0.00036489303,0.0005622485,0.5151192,0.024901535,0.28403717,0.17365082,0.000092879774],"about_ca_topic_score_codex":0.0026504667,"about_ca_topic_score_gemma":0.005543339,"teacher_disagreement_score":0.007884241,"about_ca_system_score_codex":0.0009487912,"about_ca_system_score_gemma":0.0018520647,"threshold_uncertainty_score":0.026375413},"labels":[],"label_agreement":null},{"id":"W2898905979","doi":"10.18293/seke2018-146","title":"Prioritizing Unit Testing Effort Using Software Metrics and Machine Learning Classifiers (S)","year":2018,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Innovation and Economic Development Trois Rivières","funders":"","keywords":"Regression testing; Computer science; Machine learning; Unit testing; Artificial intelligence; Non-regression testing; Naive Bayes classifier; Software quality assurance; Software quality; Software system; White-box testing; Software; Java; Software reliability testing; Software performance testing; Manual testing; Set (abstract data type); Verification and validation; Task (project management); Support vector machine; Software development; Software construction; Operating system; Programming language; Engineering","score_opus":0.06356571408007496,"score_gpt":0.27679021252861213,"score_spread":0.2132244984485372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898905979","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6636139,0.0009488036,0.3314689,0.0002023052,0.000066545486,0.0001691652,0.00027505305,0.0018412299,0.0014140389],"genre_scores_gemma":[0.9304692,0.00010169097,0.06849547,0.00003144146,0.000028545908,0.000078154684,0.000473176,0.00005981616,0.00026258218],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99394375,0.002029178,0.00073995633,0.0011165366,0.0017521824,0.00041843235],"domain_scores_gemma":[0.9712072,0.019109225,0.0033102892,0.0015593348,0.0043545393,0.00045947568],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007902442,0.0013407858,0.0013533494,0.0057222,0.00039742538,0.001578985,0.00089455687,0.0009882593,0.0004692531],"category_scores_gemma":[0.026357261,0.00030302972,0.00080792187,0.0021647762,0.0003574576,0.0017224229,0.0006693237,0.00074078795,0.00023209103],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047323696,0.00059234217,0.18467051,0.0002601312,0.0003070359,0.00016342143,0.0002906396,0.1513728,0.016506102,0.0010632296,0.0014769192,0.6428236],"study_design_scores_gemma":[0.000020253578,0.00042984236,0.037932705,0.00003685358,0.00006776627,0.00008728738,0.000094802796,0.95089585,0.008805426,0.0010798936,0.0005155362,0.000033754757],"about_ca_topic_score_codex":0.0036925548,"about_ca_topic_score_gemma":0.0043729716,"teacher_disagreement_score":0.007902442,"about_ca_system_score_codex":0.00096349313,"about_ca_system_score_gemma":0.0011208956,"threshold_uncertainty_score":0.04179263},"labels":[],"label_agreement":null},{"id":"W2899060086","doi":"10.5539/cis.v11n4p45","title":"Software Effort Estimation Risk Management over Projects Portfolio","year":2018,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Bootstrapping (finance); Project portfolio management; Flexibility (engineering); Portfolio; Risk management; Reliability (semiconductor); Risk analysis (engineering); Limiting; Software; Estimation; Project management; Power (physics); Econometrics; Statistics; Systems engineering; Finance; Mathematics","score_opus":0.009994029064184723,"score_gpt":0.2611841148351342,"score_spread":0.2511900857709495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899060086","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07042599,0.00061199523,0.9272507,0.00028008738,0.00002310135,0.000040955678,0.00006885715,0.00021192229,0.0010864377],"genre_scores_gemma":[0.90762216,0.0007171371,0.08941972,0.000054354197,0.00007600654,0.000104144354,0.0002144185,0.000078870005,0.0017132267],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.99305964,0.002504517,0.0005284836,0.0015422587,0.0019699074,0.00039522335],"domain_scores_gemma":[0.981016,0.01097553,0.0028535551,0.0020611105,0.0026584098,0.00043532686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010286981,0.0010715761,0.0012483752,0.0021568842,0.00043145157,0.002366339,0.0015378256,0.0014222187,0.00081001554],"category_scores_gemma":[0.03520974,0.00055325433,0.00081282074,0.0016940386,0.0008450208,0.003956572,0.0030883981,0.0012663455,0.00022616811],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020748022,0.00010274072,0.024524475,0.00016719829,0.00023846395,0.000254408,0.00039552647,0.7303838,0.00422983,0.052136328,0.00082295056,0.18653679],"study_design_scores_gemma":[0.0000052446944,0.000103708626,0.0040289415,0.000029070676,0.000038905757,0.000091901165,0.000065831686,0.9663105,0.0015791169,0.027030831,0.00069084746,0.000025043668],"about_ca_topic_score_codex":0.0011693109,"about_ca_topic_score_gemma":0.0006453376,"teacher_disagreement_score":0.010286981,"about_ca_system_score_codex":0.0009191459,"about_ca_system_score_gemma":0.0008908028,"threshold_uncertainty_score":0.054403424},"labels":[],"label_agreement":null},{"id":"W2899100292","doi":"10.1145/3236024.3264599","title":"WarningsGuru: integrating statistical bug models with static analysis to provide timely and specific bug warnings","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Commit; Computer science; False positive paradox; Static program analysis; Process (computing); Static analysis; Software bug; Code (set theory); Statistical model; Source code; Statistical analysis; Software; Debugging; Software engineering; Data science; Data mining; Database; Software development; Programming language; Set (abstract data type); Machine learning","score_opus":0.017186534796008346,"score_gpt":0.2631313732810692,"score_spread":0.24594483848506088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899100292","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021592587,0.0004269253,0.4889671,0.0005655175,0.0001656282,0.0002298853,0.0029711816,0.48362014,0.0014611076],"genre_scores_gemma":[0.3338537,0.0004944069,0.6232846,0.0008909013,0.0001890494,0.00058134156,0.008802945,0.02830437,0.0035986865],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99508154,0.001479752,0.00046275093,0.0010108523,0.0016877669,0.0002773896],"domain_scores_gemma":[0.9721601,0.016903099,0.0035893575,0.004785351,0.0020022606,0.00055983756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066949534,0.003156529,0.0015772284,0.008092915,0.00057025754,0.002539207,0.0028611284,0.0019253056,0.005003607],"category_scores_gemma":[0.032160573,0.0022725393,0.0026802984,0.002586685,0.0009954618,0.0054868595,0.003285226,0.0026222416,0.00306161],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012007455,0.00094821607,0.07141578,0.0016943967,0.0010659336,0.0012902861,0.0014480149,0.15801479,0.015543926,0.009070654,0.10627494,0.63203233],"study_design_scores_gemma":[0.0001269196,0.00029194108,0.0051702657,0.00015817136,0.00012257484,0.0003612258,0.00009664894,0.956712,0.008738298,0.015832921,0.0121784145,0.00021054155],"about_ca_topic_score_codex":0.0054310914,"about_ca_topic_score_gemma":0.0071605,"teacher_disagreement_score":0.008092915,"about_ca_system_score_codex":0.00090813835,"about_ca_system_score_gemma":0.0020867595,"threshold_uncertainty_score":0.03540677},"labels":[],"label_agreement":null},{"id":"W2899335321","doi":"10.1016/j.infsof.2018.10.012","title":"On semantic detection of cloud API (anti)patterns","year":2018,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Interoperability; Cloud computing; Computer science; Leverage (statistics); Application programming interface; Rest (music); Software engineering; Context (archaeology); World Wide Web; Artificial intelligence; Operating system","score_opus":0.006897650802499252,"score_gpt":0.23286842681674982,"score_spread":0.22597077601425059,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899335321","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3430284,0.001954146,0.6364846,0.000998989,0.00033289744,0.00043283092,0.0023837755,0.0045659193,0.00981842],"genre_scores_gemma":[0.73705727,0.00074130826,0.25446463,0.0003150695,0.00014486212,0.00012496604,0.0035019761,0.00021708151,0.0034327274],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9969241,0.00033562252,0.00028802908,0.000652616,0.0014369703,0.00036266714],"domain_scores_gemma":[0.99509263,0.0013863109,0.0007485947,0.0009163488,0.0016510672,0.00020506135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009968845,0.00060004345,0.000756816,0.0052422495,0.0011337169,0.0018756909,0.001168369,0.0011192801,0.0012585897],"category_scores_gemma":[0.0057759043,0.00031599283,0.0011839861,0.0035081396,0.0007069252,0.0028249088,0.0015863737,0.0008604939,0.0006281482],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007355811,0.00080814783,0.09050189,0.00082787964,0.00027684742,0.0015484862,0.0008784459,0.0096783675,0.099364296,0.04722922,0.012787073,0.7353638],"study_design_scores_gemma":[0.000075374286,0.0003852272,0.053338494,0.00022837664,0.00038217273,0.004153594,0.0015823602,0.7686958,0.070756055,0.07056127,0.029714767,0.00012651054],"about_ca_topic_score_codex":0.005278051,"about_ca_topic_score_gemma":0.008317238,"teacher_disagreement_score":0.005278051,"about_ca_system_score_codex":0.00059335545,"about_ca_system_score_gemma":0.0019494746,"threshold_uncertainty_score":0.010494649},"labels":[],"label_agreement":null},{"id":"W2899817180","doi":"10.1007/s10664-018-9665-y","title":"An empirical study of patch uplift in rapid release development pipelines","year":2018,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Geology; Channel (broadcasting); Lead (geology); Computer science; Paleontology; Telecommunications","score_opus":0.03303537718028712,"score_gpt":0.3249213934817224,"score_spread":0.2918860163014353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899817180","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9977724,0.000050621275,0.0003737557,0.00013475225,0.0000024638614,0.000022076743,0.000055244018,0.000021239743,0.0015673697],"genre_scores_gemma":[0.999032,0.00003213973,0.0002881403,0.000024260595,0.0000049496925,0.000012470908,0.00010276831,0.000008322919,0.00049496646],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9972378,0.0008858359,0.00019940197,0.00039559798,0.0008620489,0.0004193709],"domain_scores_gemma":[0.8134998,0.12250689,0.04277676,0.0068884986,0.00921692,0.00511119],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004777736,0.00024571503,0.0002572842,0.0017466173,0.0010938628,0.0016039403,0.0010264306,0.0011097135,0.0055405325],"category_scores_gemma":[0.096568696,0.0003822416,0.0002652953,0.0019686662,0.0012756481,0.0037048075,0.0013754087,0.0022367162,0.000756643],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007566218,0.0028353329,0.9227969,0.00016113442,0.00005829839,0.0007578267,0.012045006,0.0028155833,0.0019787946,0.0047385385,0.0025731812,0.048482828],"study_design_scores_gemma":[0.00007935046,0.0012288487,0.9656322,0.00008407593,0.000049186587,0.00034078956,0.014136304,0.011784343,0.00095718406,0.0018108425,0.0038609812,0.000035946232],"about_ca_topic_score_codex":0.008194193,"about_ca_topic_score_gemma":0.009634121,"teacher_disagreement_score":0.008194193,"about_ca_system_score_codex":0.0013950785,"about_ca_system_score_gemma":0.0013529642,"threshold_uncertainty_score":0.025267363},"labels":[],"label_agreement":null},{"id":"W2899979716","doi":"","title":"Questions programmers ask during software evolution tasks","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Programmer; Categorization; Task (project management); Ask price; Context (archaeology); Program comprehension; Focus (optics); Pair programming; Human–computer interaction; Code (set theory); Software; Data science; Software development; Software engineering; World Wide Web; Programming language; Software system; Artificial intelligence; Set (abstract data type)","score_opus":0.008389132803656189,"score_gpt":0.24183370224981673,"score_spread":0.23344456944616054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899979716","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9672258,0.0005795506,0.02176919,0.0030116593,0.000053157928,0.0002218954,0.00017052116,0.00041266222,0.0065556704],"genre_scores_gemma":[0.9831199,0.0005236901,0.012264194,0.0012682411,0.000051156097,0.0002744585,0.00021023804,0.00012396547,0.0021642006],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.98386616,0.011535773,0.0008348891,0.0010929629,0.0013423066,0.001327972],"domain_scores_gemma":[0.8516256,0.12972553,0.008048789,0.0025365443,0.0050619203,0.0030016236],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012745805,0.0008824803,0.00069784815,0.0015203311,0.002706377,0.0027669535,0.0012729998,0.0049134353,0.0031438011],"category_scores_gemma":[0.10136327,0.0008862882,0.0006221481,0.0009693312,0.0022683602,0.0055354508,0.0032492403,0.0024610688,0.0008459605],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048288418,0.00026050542,0.07568966,0.0006923306,0.00004273702,0.0011433563,0.8600051,0.0005940327,0.012407732,0.0023954273,0.0041591306,0.042127125],"study_design_scores_gemma":[0.0002031724,0.0015184972,0.084533,0.00083336857,0.000113871814,0.0027323163,0.8105723,0.0062696673,0.0087317005,0.008353037,0.075852424,0.00028664345],"about_ca_topic_score_codex":0.001967105,"about_ca_topic_score_gemma":0.0016366794,"teacher_disagreement_score":0.012745805,"about_ca_system_score_codex":0.0010685694,"about_ca_system_score_gemma":0.0012130857,"threshold_uncertainty_score":0.06740713},"labels":[],"label_agreement":null},{"id":"W2900654062","doi":"10.1109/icsme.2018.00078","title":"BLIMP Tracer: Integrating Build Impact Analysis with Code Review","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Deliverable; Computer science; Software engineering; Codebase; Systems engineering; Software; Source code; Code review; Software development; Static program analysis; Operating system; Engineering","score_opus":0.017891775361206137,"score_gpt":0.3303379085554981,"score_spread":0.312446133194292,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2900654062","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07082523,0.0014514975,0.619713,0.0017255116,0.00030956022,0.0031245784,0.005031611,0.2832352,0.014583842],"genre_scores_gemma":[0.27855828,0.0008850637,0.69006854,0.0007100136,0.00019694606,0.0015142234,0.008309351,0.013589337,0.006168214],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98222685,0.0063338294,0.0014994905,0.002964166,0.0064540366,0.00052174175],"domain_scores_gemma":[0.8926047,0.06842401,0.011150431,0.012058755,0.013950952,0.0018110528],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014074786,0.0018041802,0.0012950802,0.013043871,0.0010350961,0.0044790707,0.0026413037,0.0011099064,0.0046437318],"category_scores_gemma":[0.07982141,0.0014643151,0.0010597177,0.0045718914,0.001175167,0.0071271495,0.0056271693,0.0022393984,0.0028606062],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006754483,0.0008783308,0.05454926,0.0024354327,0.0005258558,0.00087754964,0.007473427,0.02148787,0.016471876,0.009561375,0.078763105,0.80630046],"study_design_scores_gemma":[0.0005250698,0.0015140357,0.06799704,0.0012502826,0.0006138869,0.0010472182,0.0037340606,0.6931838,0.038884796,0.043198805,0.1472361,0.0008150126],"about_ca_topic_score_codex":0.0122366855,"about_ca_topic_score_gemma":0.015272074,"teacher_disagreement_score":0.014074786,"about_ca_system_score_codex":0.0026489897,"about_ca_system_score_gemma":0.0049264547,"threshold_uncertainty_score":0.07443547},"labels":[],"label_agreement":null},{"id":"W2900670268","doi":"10.1109/scam.2018.00031","title":"[Engineering Paper] SCC: Automatic Classification of Code Snippets","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software engineering; Reusability; Java; Compiler; Software construction; Static program analysis; Software mining; Programming language; Software development; Software","score_opus":0.025665090259682777,"score_gpt":0.27558256137027615,"score_spread":0.24991747111059337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2900670268","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14222437,0.0021017075,0.36301148,0.003979529,0.0055746376,0.0024554713,0.17037922,0.20875537,0.10151821],"genre_scores_gemma":[0.23497884,0.0012837916,0.3581543,0.0011452101,0.0012134223,0.00071840815,0.30516377,0.011259531,0.08608283],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988275,0.000071833005,0.00012454599,0.00028277808,0.00061188685,0.00008144249],"domain_scores_gemma":[0.9956162,0.00096375175,0.0003134173,0.00064479595,0.0022742958,0.00018752305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006083594,0.0012131914,0.00052155723,0.009362818,0.00090574124,0.0022838505,0.0012055689,0.0010796783,0.034993056],"category_scores_gemma":[0.0065733106,0.00035166333,0.00076175143,0.006608061,0.00055863435,0.0013747298,0.00097315846,0.0005959453,0.0288798],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025881862,0.00012046743,0.009120892,0.00060512783,0.000074165546,0.0009449805,0.00018124584,0.0016449097,0.016767774,0.0027609754,0.38325498,0.58426565],"study_design_scores_gemma":[0.0001687034,0.00039131794,0.042839047,0.0005052569,0.00030127307,0.0042539854,0.0005329257,0.11423331,0.08053405,0.014122197,0.741902,0.0002159077],"about_ca_topic_score_codex":0.0069471733,"about_ca_topic_score_gemma":0.012425161,"teacher_disagreement_score":0.034993056,"about_ca_system_score_codex":0.00035387807,"about_ca_system_score_gemma":0.0017673924,"threshold_uncertainty_score":0.1170634},"labels":[],"label_agreement":null},{"id":"W2901077691","doi":"10.1109/icsme.2018.00009","title":"Threats of Aggregating Software Repository Data","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artifact (error); Software; Open source software; Set (abstract data type); Open source; Database; Software engineering; Operating system; Programming language","score_opus":0.05429274992367693,"score_gpt":0.3142213925941972,"score_spread":0.25992864267052024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2901077691","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.939335,0.0006581644,0.0323282,0.0046338844,0.0002323453,0.0007453593,0.006082854,0.0010182784,0.014965868],"genre_scores_gemma":[0.96891284,0.00023935214,0.021456124,0.0009218765,0.000105292296,0.000419578,0.0069603384,0.0001838056,0.00080064195],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.8780958,0.04194081,0.010695788,0.006594007,0.059404075,0.003269506],"domain_scores_gemma":[0.51936597,0.28227264,0.06702301,0.08507066,0.042431314,0.0038363354],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03967187,0.0008857741,0.0010614778,0.011323406,0.004036582,0.004836616,0.0020689256,0.0033122543,0.00069204223],"category_scores_gemma":[0.19803782,0.00062583556,0.0011218186,0.013999498,0.003914379,0.006024033,0.0072372365,0.0050663953,0.00047628517],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001901395,0.0013611342,0.69628,0.0016204477,0.0008161717,0.004650291,0.024219176,0.023717005,0.009133096,0.044046447,0.026027303,0.16622755],"study_design_scores_gemma":[0.00040120626,0.0018164093,0.46160534,0.0032009212,0.0012939257,0.021865005,0.02776928,0.15092483,0.06593815,0.08425111,0.18017167,0.0007621341],"about_ca_topic_score_codex":0.0041421033,"about_ca_topic_score_gemma":0.0030015935,"teacher_disagreement_score":0.03967187,"about_ca_system_score_codex":0.0018148662,"about_ca_system_score_gemma":0.003239421,"threshold_uncertainty_score":0.2098074},"labels":[],"label_agreement":null},{"id":"W2901220009","doi":"10.1109/icsme.2018.00086","title":"NLP2API: Query Reformulation for Code Search Using Crowdsourced Knowledge and Extra-Large Data Analytics","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Analytics; Upload; Code (set theory); Information retrieval; Vocabulary; Data science; Crowdsourcing; World Wide Web; Natural language; Web search query; Software; Search engine; Programming language; Artificial intelligence","score_opus":0.1780982632612534,"score_gpt":0.3950507227367746,"score_spread":0.21695245947552122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2901220009","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007051294,0.00020814077,0.96564895,0.0013980923,0.0001411274,0.0007419645,0.0028452524,0.018715302,0.003249802],"genre_scores_gemma":[0.109809846,0.00023302378,0.8723955,0.0009594222,0.00018395887,0.0011510975,0.009961287,0.002115847,0.003189996],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9881345,0.0047854767,0.0009767428,0.0023619027,0.0032627939,0.00047849858],"domain_scores_gemma":[0.9789652,0.010094726,0.0007755111,0.0069707334,0.0027213152,0.0004724269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008723726,0.0017550155,0.0014931052,0.004261621,0.0020369494,0.0031336856,0.0044746357,0.0021816522,0.008344076],"category_scores_gemma":[0.04270516,0.0009175609,0.0022789347,0.004514747,0.0022209932,0.007306211,0.009763247,0.0037002764,0.0051926873],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010096058,0.0010617725,0.0049394364,0.002487393,0.00040710208,0.0012834633,0.0058786427,0.056513846,0.036165964,0.08374777,0.15804684,0.6484581],"study_design_scores_gemma":[0.00022727996,0.00016859565,0.0014531408,0.00014791534,0.00010983185,0.00039185127,0.0018945255,0.7851996,0.025914716,0.11410726,0.07022704,0.00015828705],"about_ca_topic_score_codex":0.015973559,"about_ca_topic_score_gemma":0.020093769,"teacher_disagreement_score":0.015973559,"about_ca_system_score_codex":0.0019489921,"about_ca_system_score_gemma":0.004830091,"threshold_uncertainty_score":0.04613602},"labels":[],"label_agreement":null},{"id":"W2901368228","doi":"10.1109/scam.2018.00023","title":"[Research Paper] CroLSim: Cross Language Software Similarity Detector Using API Documentation","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Documentation; Software documentation; Python (programming language); Java; Application programming interface; Source code; Software; Programming language; Software engineering; World Wide Web; Software development; Software construction","score_opus":0.049071475497126926,"score_gpt":0.4025824023901418,"score_spread":0.3535109268930149,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2901368228","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4486582,0.0042310115,0.43544152,0.0015119348,0.0006464439,0.001181745,0.0204865,0.06934602,0.018496549],"genre_scores_gemma":[0.6962178,0.0009543509,0.25437117,0.0008565331,0.00014944414,0.00047227036,0.034473352,0.0013940351,0.011111074],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99803,0.0002593365,0.0002306074,0.00045533903,0.00088427035,0.00014056986],"domain_scores_gemma":[0.9970873,0.00067848625,0.0006343703,0.000502274,0.00092598697,0.00017161608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011776614,0.0010559784,0.0008001065,0.005884655,0.0005859374,0.0014333616,0.0013549067,0.0010058173,0.0025638007],"category_scores_gemma":[0.008190908,0.00035932122,0.0008089931,0.003223025,0.0004153207,0.0024479493,0.0015757701,0.0007747927,0.0024717567],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091694883,0.0005134589,0.09006206,0.00128686,0.00032950524,0.0011430794,0.0006216374,0.008636274,0.03101385,0.0058394955,0.08181402,0.7778228],"study_design_scores_gemma":[0.0002029014,0.0011984366,0.1195989,0.00031027177,0.00035930387,0.005531406,0.00088577985,0.61913973,0.102876514,0.015909657,0.13359678,0.00039028056],"about_ca_topic_score_codex":0.005508313,"about_ca_topic_score_gemma":0.0062348945,"teacher_disagreement_score":0.005884655,"about_ca_system_score_codex":0.00061672553,"about_ca_system_score_gemma":0.0011493685,"threshold_uncertainty_score":0.010952473},"labels":[],"label_agreement":null},{"id":"W2901558024","doi":"10.1109/icsme.2018.00035","title":"Predicting Higher Order Structural Feature Interactions in Variable Systems","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Feature (linguistics); Variable (mathematics); Java; Set (abstract data type); Source code; Code (set theory); Data mining; Feature model; Feature vector; Software system; Theoretical computer science; Software; Artificial intelligence; Machine learning; Mathematics; Programming language","score_opus":0.016285454132444862,"score_gpt":0.28490045105095213,"score_spread":0.26861499691850727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2901558024","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8365583,0.0004676529,0.15757264,0.00023529097,0.000023185803,0.000043904445,0.000737983,0.0026067058,0.0017543641],"genre_scores_gemma":[0.9692171,0.00006820574,0.02914521,0.000034331377,0.000014081427,0.000033058757,0.0009292143,0.00017403746,0.00038469688],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99849784,0.00029139704,0.00007827214,0.0003542609,0.0005548804,0.00022330387],"domain_scores_gemma":[0.9889352,0.0071591707,0.0019426383,0.0008736073,0.00072652835,0.00036293853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010128673,0.00077269506,0.00077024876,0.0024078474,0.000744837,0.0011794649,0.0010719146,0.0010929935,0.0018264602],"category_scores_gemma":[0.0097250035,0.0005630847,0.000936223,0.0013733892,0.0009774356,0.0018925464,0.0012310985,0.001202295,0.00040760267],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005792877,0.00041438238,0.29822916,0.00045334367,0.0002286952,0.0015450898,0.0005493615,0.5572,0.029391104,0.011205594,0.003953899,0.096250154],"study_design_scores_gemma":[0.000013577774,0.00011466976,0.02929052,0.000014066867,0.000048988604,0.0002342015,0.00011073157,0.95211756,0.005180378,0.011663479,0.001183219,0.000028461118],"about_ca_topic_score_codex":0.004023961,"about_ca_topic_score_gemma":0.0057758363,"teacher_disagreement_score":0.004023961,"about_ca_system_score_codex":0.0008449781,"about_ca_system_score_gemma":0.0008183343,"threshold_uncertainty_score":0.008001089},"labels":[],"label_agreement":null},{"id":"W2901760450","doi":"","title":"Automatic Refactoring for Renamed Clones in Test Code","year":2018,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Waterloo","keywords":"Code refactoring; Programming language; Computer science; Test (biology); Code (set theory); Null (SQL); Operating system; Software engineering; Biology; Data mining; Software","score_opus":0.016356930367944826,"score_gpt":0.24085531948425407,"score_spread":0.22449838911630923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2901760450","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22281075,0.0007657067,0.7455276,0.00034422756,0.000089869565,0.00042435125,0.00034944052,0.027757771,0.0019303416],"genre_scores_gemma":[0.3715124,0.0003747117,0.62024134,0.00028766113,0.000043525204,0.00021877616,0.0014492216,0.0032048486,0.002667478],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99624026,0.0010221208,0.000527209,0.00072315027,0.0012754305,0.00021195918],"domain_scores_gemma":[0.9781506,0.009394562,0.003720486,0.0050797635,0.00341947,0.00023509473],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028383418,0.0009801586,0.0009110156,0.002623584,0.00048915605,0.0010684432,0.0020383215,0.00089248724,0.001179215],"category_scores_gemma":[0.022016166,0.00065038836,0.0012607625,0.001371426,0.0006563331,0.0017978956,0.0012632933,0.001099029,0.0006430559],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026729266,0.000246895,0.03770739,0.0005481507,0.00013568424,0.0014467369,0.0016696665,0.017473327,0.124338165,0.0058108773,0.004347475,0.8060084],"study_design_scores_gemma":[0.00021243826,0.0006901716,0.026894478,0.00049237796,0.0004231989,0.003476057,0.00062860595,0.5111323,0.39877588,0.0115071125,0.045516092,0.0002513655],"about_ca_topic_score_codex":0.0017200068,"about_ca_topic_score_gemma":0.0031664302,"teacher_disagreement_score":0.0028383418,"about_ca_system_score_codex":0.00063800585,"about_ca_system_score_gemma":0.001411984,"threshold_uncertainty_score":0.015010774},"labels":[],"label_agreement":null},{"id":"W2901830053","doi":"10.1109/scam.2018.00025","title":"[Research Paper] On the Use of Machine Learning Techniques Towards the Design of Cloud Based Automatic Code Clone Validation Tools","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Maintainability; Software maintenance; Source code; Software system; Programming language; Software; Code (set theory); Machine learning; Software engineering; Artificial intelligence; Data mining","score_opus":0.2105035792821233,"score_gpt":0.3546240510437641,"score_spread":0.1441204717616408,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2901830053","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011359362,0.000466829,0.97893953,0.00069098675,0.00012861914,0.00031327488,0.00012323205,0.004682507,0.0032956614],"genre_scores_gemma":[0.10667152,0.00051256566,0.8871575,0.0004953593,0.00007095411,0.00018721107,0.00054412987,0.00050761865,0.0038531397],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973768,0.00058628194,0.00019990353,0.00054313097,0.0010995978,0.00019439017],"domain_scores_gemma":[0.99248874,0.002601053,0.0006838525,0.0015914222,0.0023702548,0.00026460353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032401292,0.0006803111,0.0005779705,0.0021070354,0.00080112,0.00233773,0.0021977527,0.0014842227,0.0030170614],"category_scores_gemma":[0.011176977,0.00045181855,0.001017819,0.001683623,0.0010672656,0.0029369867,0.001334636,0.0014352602,0.0017893456],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028185642,0.00041605296,0.00973898,0.00057707186,0.00013880494,0.00062309054,0.00076261867,0.04068556,0.035952687,0.038134374,0.019547628,0.85314137],"study_design_scores_gemma":[0.000092394526,0.0003603186,0.0064845355,0.00031677147,0.00013451802,0.0010103207,0.00024375865,0.79852074,0.07731147,0.026729286,0.088693365,0.00010258527],"about_ca_topic_score_codex":0.004019484,"about_ca_topic_score_gemma":0.0040430976,"teacher_disagreement_score":0.004019484,"about_ca_system_score_codex":0.0010323963,"about_ca_system_score_gemma":0.0020740496,"threshold_uncertainty_score":0.01713562},"labels":[],"label_agreement":null},{"id":"W2902192294","doi":"10.1002/smr.2149","title":"Empirical evaluation of an entropy‐based approach to estimation variation of software development effort","year":2018,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Jackknife resampling; Gaussian; Mathematics; Entropy (arrow of time); Entropy estimation; Computer science; Algorithm; Variation (astronomy); Statistics; Applied mathematics; Estimator","score_opus":0.04113808203459276,"score_gpt":0.33910532306484553,"score_spread":0.2979672410302528,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2902192294","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81381357,0.00055453577,0.18321875,0.00012901134,0.000031799744,0.000089081565,0.00031581376,0.00030828352,0.001539146],"genre_scores_gemma":[0.97139704,0.000066824214,0.02797695,0.000012112384,0.00001307608,0.000038927326,0.0003344571,0.0000151343875,0.00014549111],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916624,0.00472111,0.00050708745,0.0007354099,0.0022177997,0.000156055],"domain_scores_gemma":[0.94474316,0.045004327,0.0029154688,0.0033446949,0.003588888,0.00040343238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012779691,0.0005449945,0.00054526964,0.004541759,0.00028895307,0.0009622657,0.00087735,0.00078143534,0.0005420594],"category_scores_gemma":[0.05434649,0.00018962898,0.00047398778,0.0024935352,0.00069945544,0.0014655412,0.0014328661,0.00068695657,0.000107584645],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015190735,0.0005865266,0.18914846,0.0006419926,0.000828117,0.0002359159,0.0015022444,0.40946087,0.008637884,0.008017653,0.0012517128,0.37816942],"study_design_scores_gemma":[0.00002521792,0.0005257525,0.07184129,0.00007045618,0.00006299248,0.00018288108,0.00023596009,0.91826946,0.004966276,0.0031749841,0.000583764,0.000060975388],"about_ca_topic_score_codex":0.0017023197,"about_ca_topic_score_gemma":0.0016381306,"teacher_disagreement_score":0.012779691,"about_ca_system_score_codex":0.00079086766,"about_ca_system_score_gemma":0.0004978381,"threshold_uncertainty_score":0.0675863},"labels":[],"label_agreement":null},{"id":"W2903407100","doi":"10.4204/eptcs.284.7","title":"A Notebook Format for the Holistic Design of Embedded Systems (Tool Paper)","year":2018,"lang":"en","type":"article","venue":"Electronic Proceedings in Theoretical Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Documentation; Code (set theory); Software; Programming language; Source code; Product (mathematics); Extension (predicate logic); Software engineering; Set (abstract data type)","score_opus":0.017161912880920394,"score_gpt":0.2750405078965424,"score_spread":0.257878595015622,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2903407100","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013088614,0.00036412274,0.9590811,0.00034143543,0.00092412817,0.00040753707,0.0012479628,0.009430703,0.026894191],"genre_scores_gemma":[0.018711396,0.0009834896,0.8883785,0.0005279137,0.00044413327,0.0012827233,0.0035785558,0.00507065,0.08102266],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99853146,0.00043867092,0.00014178679,0.00018804237,0.0006123755,0.000087587796],"domain_scores_gemma":[0.99519914,0.0017312876,0.00022418259,0.0013527015,0.0011790035,0.00031364564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018332152,0.001633379,0.0006357409,0.0015087381,0.0007830529,0.0038130826,0.0023946774,0.0016970218,0.0897865],"category_scores_gemma":[0.008827106,0.0007256714,0.0007626085,0.0015441263,0.0009981361,0.002953774,0.002312316,0.002185282,0.025807073],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022785288,0.00016269568,0.00026654458,0.0012450357,0.000029279568,0.0008046735,0.0012205004,0.0112508945,0.020583076,0.19150718,0.21782507,0.5548773],"study_design_scores_gemma":[0.00004759437,0.00025347946,0.0003216242,0.00030982078,0.000011493906,0.00067565904,0.00016022053,0.0062053218,0.008056429,0.028967464,0.95493585,0.00005508831],"about_ca_topic_score_codex":0.00072442566,"about_ca_topic_score_gemma":0.0008375994,"teacher_disagreement_score":0.0897865,"about_ca_system_score_codex":0.00074383477,"about_ca_system_score_gemma":0.0011721819,"threshold_uncertainty_score":0.30036575},"labels":[],"label_agreement":null},{"id":"W2904569853","doi":"10.1007/s10515-018-0247-4","title":"Efficient elicitation of software configurations using crowd preferences and domain knowledge","year":2018,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Process (computing); Domain (mathematical analysis); Crowds; Software; ENCODE; Markov decision process; Human–computer interaction; Software engineering; Markov process","score_opus":0.019445945190041933,"score_gpt":0.2811641795886182,"score_spread":0.2617182343985763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2904569853","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46341708,0.00018031948,0.523251,0.000336218,0.00003548788,0.0004036121,0.00089907646,0.00095606974,0.010521087],"genre_scores_gemma":[0.86588836,0.00008683458,0.1313271,0.000059021364,0.000021001068,0.00025743674,0.0010110354,0.00008607687,0.0012631075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99590474,0.0017388614,0.00022139697,0.000605561,0.0012641164,0.00026528395],"domain_scores_gemma":[0.98068213,0.013864564,0.00089838,0.0022940838,0.0017639494,0.0004968174],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030077985,0.0009711707,0.0013543449,0.0019414868,0.000861576,0.0015159284,0.0011831332,0.00138869,0.0040503214],"category_scores_gemma":[0.02025258,0.0007058313,0.00067978166,0.0019422227,0.0005875909,0.0026398103,0.0028848357,0.00081800064,0.0010076339],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030569404,0.001091215,0.017645186,0.0011035395,0.00019531744,0.0006905039,0.0022552162,0.24307014,0.08473395,0.019577734,0.006753661,0.61982656],"study_design_scores_gemma":[0.00014182097,0.00036732765,0.006784063,0.00008180562,0.00005205409,0.00019295138,0.002219795,0.9034305,0.03419269,0.048864037,0.0035831518,0.00008977018],"about_ca_topic_score_codex":0.0020689506,"about_ca_topic_score_gemma":0.0049913623,"teacher_disagreement_score":0.0040503214,"about_ca_system_score_codex":0.00080035854,"about_ca_system_score_gemma":0.0015219279,"threshold_uncertainty_score":0.01590693},"labels":[],"label_agreement":null},{"id":"W2905655740","doi":"10.3390/info10010006","title":"A Comparison of Word Embeddings and N-gram Models for DBpedia Type and Invalid Entity Detection","year":2018,"lang":"en","type":"article","venue":"Information","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Entity linking; Information retrieval; Linked data; Word (group theory); Natural language processing; Named-entity recognition; Type (biology); Named entity; Cluster analysis; Artificial intelligence; Knowledge base; Task (project management); Semantic Web; Mathematics; Biology","score_opus":0.036744719496654005,"score_gpt":0.32420075310596663,"score_spread":0.2874560336093126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2905655740","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17017396,0.005215295,0.80081695,0.001136942,0.0008278907,0.00054321816,0.0036296442,0.010251203,0.0074049183],"genre_scores_gemma":[0.39599502,0.0019880794,0.5877478,0.00035129787,0.00024272136,0.00035679515,0.008871814,0.00095368165,0.0034927903],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941109,0.0028966798,0.000494813,0.0009807358,0.001270842,0.000246083],"domain_scores_gemma":[0.9793761,0.013167795,0.00078764703,0.0026511794,0.0034630504,0.00055423885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060799597,0.0024927196,0.0013548081,0.00412825,0.0009891589,0.0030626887,0.0016451495,0.0017600113,0.0022668885],"category_scores_gemma":[0.02820952,0.0006334465,0.001236627,0.0033323446,0.00062429503,0.008386907,0.0022504602,0.0024549272,0.0028718275],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027947726,0.0012482534,0.018148886,0.0010778555,0.001032575,0.0002738325,0.0009734683,0.057549585,0.011619866,0.01079155,0.015768593,0.8787207],"study_design_scores_gemma":[0.000081542916,0.00044109815,0.003616187,0.00011937993,0.00015464005,0.0002692576,0.0005964507,0.9691075,0.008360748,0.010977169,0.0061576217,0.0001184995],"about_ca_topic_score_codex":0.010489171,"about_ca_topic_score_gemma":0.013654749,"teacher_disagreement_score":0.010489171,"about_ca_system_score_codex":0.0011184823,"about_ca_system_score_gemma":0.0020424898,"threshold_uncertainty_score":0.032154262},"labels":[],"label_agreement":null},{"id":"W2907438654","doi":"","title":"Revisiting \"Programmers' Build Errors\" in the Visual Studio Context","year":2018,"lang":"en","type":"article","venue":"Mining Software Repositories","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Deliverable; Computer science; TRACE (psycholinguistics); Context (archaeology); Replicate; Workspace; Studio; Codebase; Code (set theory); Software engineering; Human–computer interaction; Programming language; Source code; World Wide Web; Artificial intelligence; Systems engineering","score_opus":0.01943302576174965,"score_gpt":0.3028428703419978,"score_spread":0.28340984458024815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2907438654","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98104846,0.0013913554,0.0101268925,0.001208826,0.00007424561,0.00007540589,0.0011333587,0.0003172558,0.0046241465],"genre_scores_gemma":[0.9875158,0.0004568835,0.008582102,0.0003561976,0.000086058746,0.0000687632,0.0016066785,0.00026204082,0.001065485],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9809806,0.006800051,0.0019134311,0.0034489897,0.0056789354,0.0011780631],"domain_scores_gemma":[0.8261028,0.100335665,0.041094806,0.015728327,0.014734277,0.002004064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013781794,0.0006318505,0.0005083072,0.009184364,0.0014392821,0.0039651375,0.0013300172,0.0009808649,0.0010321592],"category_scores_gemma":[0.11348235,0.00049982965,0.00060740515,0.009544346,0.001914285,0.0060046,0.0038991577,0.0019211192,0.000394993],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001772377,0.00017115835,0.87633836,0.00052997086,0.00013585619,0.0007265903,0.0275237,0.0009882682,0.0015180247,0.0025084605,0.0053024734,0.08407994],"study_design_scores_gemma":[0.00002782595,0.00021838475,0.9237546,0.00053379586,0.00016534458,0.001436999,0.032774717,0.007738865,0.0025405379,0.006870514,0.02386049,0.00007794557],"about_ca_topic_score_codex":0.010240169,"about_ca_topic_score_gemma":0.024988089,"teacher_disagreement_score":0.013781794,"about_ca_system_score_codex":0.0010742162,"about_ca_system_score_gemma":0.001840143,"threshold_uncertainty_score":0.07288599},"labels":[],"label_agreement":null},{"id":"W2909172538","doi":"10.1109/tse.2019.2891758","title":"The Impact of Correlated Metrics on the Interpretation of Defect Models","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Australian Research Council","keywords":"Consistency (knowledge bases); Ranking (information retrieval); Interpretation (philosophy); Computer science; Metric (unit); Data mining; Statistics; Machine learning; Artificial intelligence; Mathematics; Programming language","score_opus":0.014277181931630203,"score_gpt":0.2477259451862916,"score_spread":0.2334487632546614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2909172538","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49013096,0.0030791068,0.49266303,0.0017900303,0.00030642765,0.00030613018,0.002049135,0.0029730906,0.0067021204],"genre_scores_gemma":[0.86897737,0.00038896248,0.12665683,0.00027228965,0.00008461954,0.00011606374,0.002318172,0.00077359116,0.00041193303],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.92465204,0.041829336,0.005449464,0.008062754,0.018974144,0.001032183],"domain_scores_gemma":[0.5389713,0.33771,0.037503302,0.057263017,0.027108101,0.0014443466],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.043280363,0.0020499502,0.001699579,0.006275942,0.0010460936,0.005392224,0.0018754132,0.0013643613,0.0010323805],"category_scores_gemma":[0.28535756,0.000909962,0.0022966275,0.0046757692,0.0026265462,0.0056612906,0.0032448731,0.0033409488,0.00047136802],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013962414,0.0004236394,0.34005594,0.0017917692,0.002403739,0.001022628,0.0032245712,0.23079541,0.015635068,0.026094012,0.007756543,0.36940044],"study_design_scores_gemma":[0.0002283128,0.0015514109,0.15296496,0.0009763682,0.0012094968,0.0018176361,0.0025238264,0.67530805,0.026802607,0.122178644,0.013958805,0.0004798193],"about_ca_topic_score_codex":0.0032307904,"about_ca_topic_score_gemma":0.004661975,"teacher_disagreement_score":0.95671964,"about_ca_system_score_codex":0.0022847326,"about_ca_system_score_gemma":0.0026562195,"threshold_uncertainty_score":0.2288912},"labels":[],"label_agreement":null},{"id":"W2909755202","doi":"10.1109/tse.2019.2893171","title":"Too Many User-Reviews! What Should App Developers Look at First?","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Polytechnique Montréal","keywords":"Key (lock); Computer science; Android (operating system); Mobile apps; World Wide Web; App store; Set (abstract data type); Android app; Star (game theory); Mobile device; Multimedia; Computer security","score_opus":0.023871746814439637,"score_gpt":0.2479435287887195,"score_spread":0.22407178197427988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2909755202","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9583822,0.0064423224,0.00727736,0.007325971,0.00049711886,0.00044776892,0.0032061024,0.0008745008,0.015546667],"genre_scores_gemma":[0.9819603,0.0014633794,0.0058118463,0.0021397753,0.00028310524,0.00020815413,0.0017178777,0.0001727244,0.006242947],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99455875,0.0019585474,0.0004797781,0.00065255136,0.002060124,0.0002902281],"domain_scores_gemma":[0.9491192,0.025369791,0.0082098525,0.0024736875,0.012648061,0.0021794455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003156778,0.0005860475,0.00068723987,0.0014337867,0.0010585921,0.0023247688,0.0004966543,0.001014656,0.0029581166],"category_scores_gemma":[0.04923952,0.00042334464,0.0005049134,0.0016493676,0.00046865857,0.0031814084,0.0007532365,0.0009158319,0.0028432868],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007060045,0.0004078387,0.6072073,0.0016782453,0.00033071468,0.0012732908,0.0168376,0.00045560522,0.008687019,0.0008590624,0.087299794,0.27425745],"study_design_scores_gemma":[0.000065510896,0.0007639156,0.8512921,0.0005927401,0.00026709528,0.0031110824,0.024289219,0.0051480993,0.004842173,0.0015103079,0.107933864,0.0001839892],"about_ca_topic_score_codex":0.0050200587,"about_ca_topic_score_gemma":0.011518411,"teacher_disagreement_score":0.0050200587,"about_ca_system_score_codex":0.0007288515,"about_ca_system_score_gemma":0.000665265,"threshold_uncertainty_score":0.016694844},"labels":[],"label_agreement":null},{"id":"W2909995151","doi":"10.1002/smr.2151","title":"A measurement framework for software product maturity assessment","year":2019,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Capability Maturity Model; Computer science; Software quality control; Product (mathematics); Quality (philosophy); Maturity (psychological); Process management; Software quality; Software measurement; Software quality analyst; Software; Systems engineering; Software engineering; Software development; Engineering","score_opus":0.022701644488524977,"score_gpt":0.30686463390930613,"score_spread":0.28416298942078116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2909995151","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006752675,0.00084193464,0.9845595,0.0004968473,0.000053532724,0.00070745393,0.00035906408,0.0007730342,0.0054559917],"genre_scores_gemma":[0.11625876,0.00045771254,0.88057256,0.00008560192,0.000034998,0.0014263171,0.0006302846,0.00005963856,0.00047418906],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9647888,0.01975815,0.004319761,0.0022294107,0.008190297,0.00071359787],"domain_scores_gemma":[0.9457446,0.025147153,0.00758794,0.003630403,0.016952997,0.0009369511],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032762658,0.0018966871,0.00092767214,0.014541105,0.0014316665,0.005657373,0.0025059506,0.0022108918,0.0020128493],"category_scores_gemma":[0.063829154,0.0007645098,0.0017209091,0.009446841,0.002304554,0.00884478,0.0034620822,0.0034482237,0.00078928284],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000104004284,0.00053120626,0.018402213,0.001366255,0.00024670942,0.00018813537,0.002826406,0.032581847,0.0076735,0.5669774,0.007214491,0.3618878],"study_design_scores_gemma":[0.000121829005,0.0013405966,0.02857292,0.002779136,0.00026475766,0.000816023,0.002947565,0.47807086,0.011517334,0.3946361,0.0784849,0.00044805033],"about_ca_topic_score_codex":0.0066107633,"about_ca_topic_score_gemma":0.0036271974,"teacher_disagreement_score":0.032762658,"about_ca_system_score_codex":0.0049845157,"about_ca_system_score_gemma":0.005792959,"threshold_uncertainty_score":0.1732676},"labels":[],"label_agreement":null},{"id":"W2911055565","doi":"10.1109/ieem.2018.8607699","title":"Towards a Knowledge based Support for Risk Engineering When Elaborating Offer in Response to a Customer Demand","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Agence Nationale de la Recherche","keywords":"Computer science; Risk analysis (engineering); Knowledge base; Risk management; Knowledge engineering; Call for bids; Knowledge management; Artificial intelligence; Business; Procurement","score_opus":0.01805150935649408,"score_gpt":0.29473170812374666,"score_spread":0.2766801987672526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911055565","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012721908,0.00024067744,0.97666025,0.0013611433,0.000034035234,0.00025918,0.00042553523,0.0038288562,0.0044683428],"genre_scores_gemma":[0.12849545,0.0005293438,0.86483586,0.00039328373,0.00010263706,0.00018766303,0.0017099562,0.00037780628,0.0033679418],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9943566,0.0016887406,0.0006110103,0.0010615322,0.00197417,0.00030790377],"domain_scores_gemma":[0.9849569,0.009035627,0.0008830405,0.0029149107,0.0016870883,0.0005224764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056251953,0.0010854949,0.0010680951,0.0045739403,0.0011790526,0.008137892,0.0037273315,0.003843469,0.007127841],"category_scores_gemma":[0.026302943,0.000888292,0.0018666231,0.0025814432,0.001304524,0.00998095,0.004485771,0.0031567765,0.0040796995],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078150403,0.0016001867,0.0050611794,0.0019303095,0.0003181087,0.003445789,0.010725031,0.0370901,0.046402887,0.11668274,0.016769532,0.7591926],"study_design_scores_gemma":[0.0002887453,0.00041476058,0.004020364,0.001017737,0.0005876481,0.0036749695,0.0049745315,0.5111746,0.06171747,0.23341419,0.17840019,0.0003147827],"about_ca_topic_score_codex":0.0035760065,"about_ca_topic_score_gemma":0.0025185235,"teacher_disagreement_score":0.008137892,"about_ca_system_score_codex":0.0010811528,"about_ca_system_score_gemma":0.0020933663,"threshold_uncertainty_score":0.029749215},"labels":[],"label_agreement":null},{"id":"W2911278647","doi":"","title":"2nd workshop on DevOps and software analytics for continuous engineering and improvement","year":2018,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; École de Technologie Supérieure; University of Victoria","funders":"","keywords":"DevOps; Software deployment; Software engineering; Computer science; IBM; Toolchain; Software analytics; Software development; Software; Analytics; Software system; Data science; Software construction; Operating system","score_opus":0.014449569783324866,"score_gpt":0.2449149988382957,"score_spread":0.23046542905497083,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911278647","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.050971236,0.02105455,0.6089954,0.17961894,0.054608326,0.0013923466,0.002792167,0.0024506224,0.078116454],"genre_scores_gemma":[0.3287907,0.023330925,0.33177865,0.018418733,0.024343992,0.0030222745,0.007655746,0.0028068696,0.25985214],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9937045,0.0027706032,0.00028432024,0.0011331837,0.0012509541,0.0008564417],"domain_scores_gemma":[0.9867916,0.0068435427,0.00027631473,0.0010000973,0.002572052,0.0025163842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018652273,0.0018244773,0.0010439397,0.0017243795,0.0015505643,0.011856957,0.0030850174,0.0042508976,0.017539391],"category_scores_gemma":[0.016853997,0.0009015725,0.0021208185,0.0015703727,0.001971068,0.008217291,0.0071328096,0.0073518376,0.003560382],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007774414,0.0009617852,0.0021114002,0.00091333914,0.00018358848,0.00091340236,0.012084475,0.012312454,0.0060761836,0.12307794,0.45234925,0.3882388],"study_design_scores_gemma":[0.00009983132,0.00034450035,0.0013143184,0.00080587424,0.000051243904,0.000298382,0.00465052,0.014695445,0.0035307228,0.06331452,0.91078484,0.000109799796],"about_ca_topic_score_codex":0.002053299,"about_ca_topic_score_gemma":0.0024157327,"teacher_disagreement_score":0.018652273,"about_ca_system_score_codex":0.003061188,"about_ca_system_score_gemma":0.004571305,"threshold_uncertainty_score":0.09864384},"labels":[],"label_agreement":null},{"id":"W2911476493","doi":"10.1007/s11219-018-9428-4","title":"API trustworthiness: an ontological approach for software library adoption","year":2019,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Trustworthiness; License; Software engineering; Quality (philosophy); Application programming interface; Software; World Wide Web; Data science; Computer security","score_opus":0.046976021699527935,"score_gpt":0.3156316965895039,"score_spread":0.268655674889976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911476493","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08397589,0.0005239139,0.8737751,0.004817682,0.00013675525,0.00034436182,0.0004238414,0.00058602844,0.035416357],"genre_scores_gemma":[0.7668827,0.00061476615,0.22703916,0.00038642174,0.000120280056,0.00028098884,0.0005863461,0.00022924795,0.0038599654],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99239606,0.002106194,0.0011136475,0.0009717308,0.0027314115,0.0006809585],"domain_scores_gemma":[0.98216337,0.0067702467,0.0020743306,0.004167372,0.0040921303,0.00073250534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00840191,0.00060626515,0.0006520157,0.009464205,0.0037658429,0.007494507,0.0022829508,0.0020519071,0.002335867],"category_scores_gemma":[0.025968952,0.0009236011,0.002658297,0.00679422,0.00607375,0.017390586,0.0053547076,0.003571856,0.00053840346],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034934135,0.00012518845,0.0133439265,0.00026148482,0.00010948286,0.00043625198,0.0092762485,0.0031775306,0.0028287505,0.9125018,0.0014622732,0.05644218],"study_design_scores_gemma":[0.00002443245,0.00010720567,0.015460284,0.0007016717,0.00053392025,0.0010921424,0.0122533,0.087410025,0.0052445866,0.8193418,0.05768638,0.00014423183],"about_ca_topic_score_codex":0.019207094,"about_ca_topic_score_gemma":0.015813293,"teacher_disagreement_score":0.019207094,"about_ca_system_score_codex":0.0046662055,"about_ca_system_score_gemma":0.0056302273,"threshold_uncertainty_score":0.04443413},"labels":[],"label_agreement":null},{"id":"W2911710442","doi":"","title":"Proceedings of the VARiability for You Workshop: Variability Modeling Made Useful for Everyone","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Data science","score_opus":0.040392155909693486,"score_gpt":0.2806550905942591,"score_spread":0.2402629346845656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911710442","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041000735,0.008937594,0.67881423,0.11743608,0.070658244,0.00066864863,0.0021017473,0.0068139103,0.073568806],"genre_scores_gemma":[0.2937063,0.007615725,0.28802407,0.01938831,0.020424662,0.0010800421,0.0067577967,0.009251915,0.35375115],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9964,0.0014480184,0.00013552496,0.0004955345,0.0011111028,0.000409793],"domain_scores_gemma":[0.985688,0.0038418924,0.0004520677,0.0019121655,0.004554765,0.0035510843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013432027,0.0014309612,0.0010969192,0.0006221116,0.0021155162,0.006797255,0.0025839012,0.0034497858,0.025520219],"category_scores_gemma":[0.015723933,0.0008824129,0.0016328256,0.00069668004,0.0010017677,0.00768332,0.0073446403,0.0076225917,0.008837346],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005623775,0.00036800496,0.001350017,0.00029210452,0.00009169457,0.00053249206,0.0023187958,0.003362524,0.0074336887,0.019429542,0.7746618,0.18959694],"study_design_scores_gemma":[0.00012562229,0.00034846857,0.0018659991,0.0004256403,0.00015077808,0.00048527383,0.0021849729,0.013490356,0.0091714645,0.04983098,0.92172146,0.00019914078],"about_ca_topic_score_codex":0.0020027165,"about_ca_topic_score_gemma":0.005547073,"teacher_disagreement_score":0.025520219,"about_ca_system_score_codex":0.0010881678,"about_ca_system_score_gemma":0.0035984495,"threshold_uncertainty_score":0.08537358},"labels":[],"label_agreement":null},{"id":"W2911789761","doi":"10.1007/s10664-018-9671-0","title":"Automatic query reformulation for code search using crowdsourced knowledge","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Information retrieval; Web search query; Java; Code (set theory); Query expansion; Search engine; Web query classification; XPath; Precision and recall; Natural language; Programming language; World Wide Web; Artificial intelligence; XML; Set (abstract data type)","score_opus":0.0489429157531969,"score_gpt":0.33474163569030263,"score_spread":0.2857987199371057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911789761","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08637204,0.0021826508,0.8493398,0.002100178,0.00050653814,0.0017377636,0.009045261,0.034911204,0.013804503],"genre_scores_gemma":[0.42132717,0.0007229435,0.5393496,0.0006939675,0.00028367955,0.0009966008,0.025087005,0.002061641,0.009477247],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99297935,0.0018670734,0.00051925133,0.0012200715,0.0028828748,0.00053133897],"domain_scores_gemma":[0.98790026,0.005119511,0.00051581755,0.002669304,0.0034288252,0.00036633792],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025995432,0.0014978731,0.0018692572,0.007863757,0.0021715374,0.0026772642,0.002772435,0.002239003,0.014339402],"category_scores_gemma":[0.023804503,0.0005852267,0.0016973112,0.004772567,0.001071148,0.004800716,0.005934762,0.0016621961,0.006777961],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010769182,0.00079614477,0.0033213915,0.0017116144,0.00019550504,0.0007083854,0.0022844686,0.018321684,0.05665681,0.016914438,0.11044601,0.7875666],"study_design_scores_gemma":[0.0003655837,0.0004220522,0.004428033,0.0003410807,0.00031658972,0.0007783648,0.0038087082,0.7850179,0.06859254,0.054328077,0.081371166,0.00022982674],"about_ca_topic_score_codex":0.014811992,"about_ca_topic_score_gemma":0.020216096,"teacher_disagreement_score":0.014811992,"about_ca_system_score_codex":0.0018420956,"about_ca_system_score_gemma":0.004259887,"threshold_uncertainty_score":0.047970116},"labels":[],"label_agreement":null},{"id":"W2912212393","doi":"","title":"Proceedings of the 2nd International Workshop on Recommendation Systems for Software Engineering","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; McGill University","funders":"","keywords":"Computer science; Code refactoring; Usability; Debugging; Software engineering; Leverage (statistics); Recommender system; Variety (cybernetics); World Wide Web; Software; Data science; Human–computer interaction","score_opus":0.017101609198645056,"score_gpt":0.2610871166481692,"score_spread":0.24398550744952413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2912212393","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014777097,0.07059479,0.81062454,0.029352132,0.019886099,0.0017529586,0.0022097398,0.0055355537,0.045267135],"genre_scores_gemma":[0.08698664,0.05290101,0.68004274,0.00757329,0.017113324,0.0023472684,0.013127484,0.0020009505,0.13790733],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98715067,0.0062057134,0.0013931352,0.0019171629,0.0026961057,0.0006371869],"domain_scores_gemma":[0.9786881,0.010040538,0.00045649603,0.003953246,0.0059599807,0.0009015111],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021012893,0.003016528,0.0048379335,0.00361387,0.0020673177,0.008661271,0.0046861973,0.0058835475,0.035547893],"category_scores_gemma":[0.034292966,0.0018207809,0.0031450791,0.003540304,0.002045575,0.014349597,0.0032292136,0.0066990484,0.015245034],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008644347,0.0007104389,0.0025050056,0.0011460694,0.0007169327,0.00039809922,0.0010720282,0.0069167404,0.0037982222,0.035868954,0.30859753,0.6374055],"study_design_scores_gemma":[0.000334646,0.0006063645,0.004080662,0.0010950364,0.00043090622,0.00067018974,0.000655637,0.06969768,0.0032611815,0.051639386,0.86728805,0.00024038061],"about_ca_topic_score_codex":0.012574003,"about_ca_topic_score_gemma":0.012055016,"teacher_disagreement_score":0.035547893,"about_ca_system_score_codex":0.0027081924,"about_ca_system_score_gemma":0.0027708127,"threshold_uncertainty_score":0.11891949},"labels":[],"label_agreement":null},{"id":"W2912360270","doi":"","title":"Proceedings of the 27th ACM SIGPLAN Conference on Programming Language Design and Implementation","year":2006,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Special Interest Group; Library science; Seniority; Operations research; Engineering","score_opus":0.026533720172754873,"score_gpt":0.29920071389131697,"score_spread":0.2726669937185621,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2912360270","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007480958,0.033778735,0.65392303,0.012081145,0.021211015,0.0022823184,0.003942438,0.016423523,0.24887687],"genre_scores_gemma":[0.03401599,0.038145788,0.52921844,0.0045421473,0.00408133,0.0030416287,0.01648872,0.008343109,0.3621229],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9943693,0.0023582124,0.00054779876,0.00070719427,0.0016947091,0.00032279931],"domain_scores_gemma":[0.9943001,0.0023224049,0.00018764345,0.0011377749,0.0014010936,0.00065088307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074613458,0.0029124955,0.0017055145,0.0015758683,0.0014867958,0.0075389594,0.0024329498,0.0018267054,0.12745921],"category_scores_gemma":[0.012771627,0.0020629477,0.0018499424,0.0013235095,0.0019536829,0.0053794906,0.0035927268,0.0061011836,0.048116513],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023880898,0.00033382513,0.0010734916,0.0008737711,0.000136437,0.00019667846,0.000654831,0.0035041957,0.0028410458,0.032325424,0.5505031,0.40731838],"study_design_scores_gemma":[0.00006472915,0.0000817417,0.00035934395,0.00039131273,0.000050964452,0.00018733484,0.00014716538,0.0052020857,0.0010023273,0.0126331635,0.97984785,0.000031965912],"about_ca_topic_score_codex":0.0066915927,"about_ca_topic_score_gemma":0.007223173,"teacher_disagreement_score":0.12745921,"about_ca_system_score_codex":0.0015793707,"about_ca_system_score_gemma":0.0069866977,"threshold_uncertainty_score":0.42639357},"labels":[],"label_agreement":null},{"id":"W2913441567","doi":"","title":"Proceedings of the 4th International Workshop on Managing Technical Debt","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Debt; Documentation; Software deployment; Software; Computer science; Software engineering; Engineering; Software development; Business; Finance; Operating system","score_opus":0.02349287177094463,"score_gpt":0.26831528191792914,"score_spread":0.2448224101469845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2913441567","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065127667,0.05628457,0.12084498,0.098088495,0.079166204,0.0006064224,0.0016856317,0.0040856255,0.63272536],"genre_scores_gemma":[0.039079264,0.026518939,0.05638709,0.009430395,0.01142708,0.0006052164,0.0054748105,0.002685779,0.84839135],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9964714,0.0010026621,0.00022557456,0.000543117,0.0012793528,0.0004778767],"domain_scores_gemma":[0.9939757,0.0013791346,0.00021636383,0.0009564664,0.0019693924,0.0015029723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058332207,0.0011908358,0.00085565116,0.0017280739,0.0023663987,0.009527105,0.0025373118,0.0030172162,0.12826698],"category_scores_gemma":[0.009790239,0.0006932261,0.0010047508,0.0019048153,0.0015179795,0.009579846,0.006240286,0.0052804975,0.051893655],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000099681,0.00012158992,0.00030811472,0.00025714037,0.00001735642,0.00012623605,0.00082938874,0.00040016527,0.0014008968,0.022272883,0.74768215,0.2264845],"study_design_scores_gemma":[0.0000073624037,0.000015704947,0.0001936339,0.00015683718,0.0000065542927,0.00007322124,0.00028914746,0.000298423,0.00029086994,0.005508047,0.9931497,0.000010494186],"about_ca_topic_score_codex":0.003061469,"about_ca_topic_score_gemma":0.0061865044,"teacher_disagreement_score":0.12826698,"about_ca_system_score_codex":0.0020377617,"about_ca_system_score_gemma":0.0035821625,"threshold_uncertainty_score":0.4290958},"labels":[],"label_agreement":null},{"id":"W2913576447","doi":"10.1109/tifs.2019.2895963","title":"Large-Scale Empirical Study of Important Features Indicative of Discovered Vulnerabilities to Assess Application Security","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Information Forensics and Security","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Feature selection; Vulnerability (computing); Secure coding; Scale (ratio); Empirical research; Machine learning; Feature (linguistics); Data science; Predictive power; Focus (optics); Selection (genetic algorithm); Data mining; Artificial intelligence; Computer security; Information security; Software security assurance; Statistics","score_opus":0.012251667946967618,"score_gpt":0.28365874951568815,"score_spread":0.2714070815687205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2913576447","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98505634,0.00036853593,0.013033171,0.00011967188,0.000009682165,0.000055485132,0.0006427576,0.00016396581,0.00055047806],"genre_scores_gemma":[0.99330807,0.00008880049,0.005722684,0.00001251048,0.000008226911,0.000030313715,0.0007384038,0.000013648068,0.000077305944],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9979486,0.000657369,0.00019512662,0.000357608,0.00070448965,0.00013689128],"domain_scores_gemma":[0.947002,0.040084425,0.0062226425,0.0035409746,0.0026063223,0.0005435942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037871713,0.0009069342,0.0004860619,0.00507532,0.00045007013,0.00070562714,0.00056823646,0.0005140341,0.00058595586],"category_scores_gemma":[0.029216744,0.0001861049,0.00062099734,0.0034910783,0.0006878462,0.0015806381,0.000824908,0.001016292,0.00020202962],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020994425,0.00065670337,0.87102896,0.00022842956,0.00035184104,0.0002859193,0.00033727344,0.033599697,0.004430741,0.00063074613,0.0017450323,0.086494766],"study_design_scores_gemma":[0.00002184234,0.0003841137,0.68555236,0.000065010485,0.0001940688,0.00067296054,0.00050334685,0.3008496,0.007686744,0.0027212128,0.0012991357,0.00004957821],"about_ca_topic_score_codex":0.0015093489,"about_ca_topic_score_gemma":0.0023626585,"teacher_disagreement_score":0.00507532,"about_ca_system_score_codex":0.00043995932,"about_ca_system_score_gemma":0.00047375384,"threshold_uncertainty_score":0.02002871},"labels":[],"label_agreement":null},{"id":"W2914035201","doi":"10.1007/s10664-019-09684-y","title":"Towards prioritizing user-related issue reports of mobile applications","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Prioritization; Computer science; Android (operating system); World Wide Web; Empirical research; Mobile apps; Process (computing); Matching (statistics); Data science; Internet privacy; Engineering; Process management","score_opus":0.00985313260491852,"score_gpt":0.2752562160281317,"score_spread":0.26540308342321317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914035201","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6087324,0.005301878,0.34316728,0.00396585,0.0006752288,0.0024356106,0.0060716546,0.0068027093,0.022847261],"genre_scores_gemma":[0.7348778,0.0012635555,0.2522573,0.00034182565,0.0003490186,0.0003812512,0.005472346,0.00036714113,0.0046897647],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98289675,0.0054452815,0.0026429805,0.0013962132,0.006555807,0.0010629455],"domain_scores_gemma":[0.8228139,0.08433201,0.02644898,0.006084991,0.055697855,0.0046221744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018070482,0.0017968118,0.0010930492,0.025633285,0.001437125,0.009897017,0.0017515486,0.0016781728,0.0026060964],"category_scores_gemma":[0.11961482,0.00068792165,0.0009395056,0.007549374,0.00047147562,0.005174757,0.0023112493,0.0016938291,0.0018675535],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001025759,0.00080135395,0.32162836,0.001772918,0.0002627872,0.00052878197,0.0030372099,0.0040267026,0.03455808,0.0054752086,0.0101180095,0.6167649],"study_design_scores_gemma":[0.00024173694,0.0022032103,0.50598246,0.0018967872,0.0015444782,0.0023143368,0.021367233,0.25235182,0.116292775,0.029470779,0.065905444,0.00042895178],"about_ca_topic_score_codex":0.008769181,"about_ca_topic_score_gemma":0.016055726,"teacher_disagreement_score":0.025633285,"about_ca_system_score_codex":0.001313951,"about_ca_system_score_gemma":0.0057641147,"threshold_uncertainty_score":0.09556693},"labels":[],"label_agreement":null},{"id":"W2914275570","doi":"","title":"Proceedings of the 9th ACM SIGPLAN-SIGSOFT workshop on Program analysis for software tools and engineering","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Constructive; Software engineering; Computer science; Software; Library science; Engineering management; Engineering; Programming language","score_opus":0.023748676904518012,"score_gpt":0.27615872877818787,"score_spread":0.2524100518736698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914275570","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009975191,0.022716982,0.82645726,0.011870225,0.01629796,0.0009651649,0.0033038286,0.011471333,0.09694212],"genre_scores_gemma":[0.055531505,0.035375316,0.6882375,0.0027204887,0.006632115,0.0012083205,0.014416021,0.010248421,0.1856304],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934683,0.0024643787,0.00067718956,0.0007810171,0.0023100202,0.00029912146],"domain_scores_gemma":[0.98208433,0.008905138,0.0005112603,0.0032761467,0.0038954902,0.0013276014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00976541,0.002076053,0.0020151974,0.0025900344,0.0013808339,0.008764395,0.0022154546,0.001828104,0.050361183],"category_scores_gemma":[0.02161019,0.0017370436,0.0018373705,0.0022665479,0.0024864753,0.0059968214,0.0029385944,0.006063126,0.02369728],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003837392,0.00047136538,0.0022071425,0.00070896663,0.00016528596,0.00027060465,0.0008440713,0.0054090573,0.003898969,0.02490862,0.41762155,0.5431106],"study_design_scores_gemma":[0.000079791665,0.00019675464,0.0016271005,0.0006706658,0.00012105989,0.00047241087,0.0003820719,0.012359414,0.0030818677,0.036765635,0.9441743,0.00006887123],"about_ca_topic_score_codex":0.006988229,"about_ca_topic_score_gemma":0.010608533,"teacher_disagreement_score":0.050361183,"about_ca_system_score_codex":0.0019913588,"about_ca_system_score_gemma":0.0075534903,"threshold_uncertainty_score":0.16847497},"labels":[],"label_agreement":null},{"id":"W2914445998","doi":"10.7763/lnse.2016.v4.234","title":"Applying Variant Variable Regularized Logistic Regression for Modeling Software Defect Predictor","year":2015,"lang":"en","type":"article","venue":"Lecture Notes on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Logistic regression; Statistics; Software; Variable (mathematics); Computer science; Regression analysis; Artificial intelligence; Econometrics; Mathematics","score_opus":0.0403115085901535,"score_gpt":0.26916388672880115,"score_spread":0.22885237813864764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914445998","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14522222,0.0016249052,0.8461419,0.000778731,0.0000936531,0.00014418799,0.000819341,0.0040299986,0.0011451426],"genre_scores_gemma":[0.8079776,0.00066265994,0.18633416,0.00019724283,0.000127973,0.00020375491,0.0020204023,0.00026502996,0.002211112],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954438,0.0023409976,0.0002590706,0.0011238853,0.0005685322,0.0002637795],"domain_scores_gemma":[0.98626685,0.009563469,0.001320698,0.0010800944,0.0016150102,0.0001538554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069533605,0.0014447964,0.001299365,0.0032204022,0.0004151353,0.001645379,0.001879672,0.0011915563,0.001167072],"category_scores_gemma":[0.022443565,0.0005031487,0.0018785712,0.0026546977,0.0005802496,0.002066999,0.0009946242,0.0022279953,0.0012333427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045498155,0.00031784485,0.08660833,0.00031065228,0.0007518775,0.00044569402,0.00030113012,0.62912285,0.00367679,0.0057456866,0.0060185855,0.26624563],"study_design_scores_gemma":[0.000012021882,0.00008334516,0.0030757545,0.000025042973,0.00005354462,0.00010862186,0.000038043196,0.9906466,0.0008552657,0.0040658396,0.0010107321,0.000025246129],"about_ca_topic_score_codex":0.007291238,"about_ca_topic_score_gemma":0.0055934805,"teacher_disagreement_score":0.007291238,"about_ca_system_score_codex":0.00078189344,"about_ca_system_score_gemma":0.0013172921,"threshold_uncertainty_score":0.036773324},"labels":[],"label_agreement":null},{"id":"W2914489208","doi":"10.1109/tse.2019.2897300","title":"Which Commits Can Be CI Skipped?","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Commit; Computer science; Process (computing); Java; Software engineering; Database; Programming language","score_opus":0.012768036516850528,"score_gpt":0.2325881247438001,"score_spread":0.2198200882269496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914489208","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7270205,0.0029642438,0.23757413,0.0017067979,0.00072547205,0.00081303116,0.00404068,0.015859265,0.009295912],"genre_scores_gemma":[0.84151417,0.0009763623,0.14353952,0.00063253316,0.00014831728,0.00022850042,0.0071938736,0.0012145743,0.0045522805],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9893652,0.000993974,0.001331597,0.0027818398,0.0047090845,0.0008183306],"domain_scores_gemma":[0.93148464,0.028503474,0.013772416,0.009772255,0.014313019,0.002154173],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007260547,0.0011997609,0.0009675661,0.003972981,0.0011746483,0.0026164425,0.0018220987,0.0011751266,0.0011160908],"category_scores_gemma":[0.059734862,0.0007372037,0.0011117825,0.0022193803,0.00077392685,0.0025487829,0.0015480474,0.0021779588,0.0009600546],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082930655,0.00034787395,0.47335035,0.0013355545,0.0003152367,0.0031534843,0.0024392908,0.019208593,0.036737606,0.0055555445,0.015528564,0.4411986],"study_design_scores_gemma":[0.00016065444,0.0014601502,0.30635396,0.0012813455,0.00089841284,0.0080000935,0.0055333306,0.46693006,0.10454127,0.027518233,0.07687652,0.0004459143],"about_ca_topic_score_codex":0.010920066,"about_ca_topic_score_gemma":0.01924554,"teacher_disagreement_score":0.010920066,"about_ca_system_score_codex":0.0009759408,"about_ca_system_score_gemma":0.0039979285,"threshold_uncertainty_score":0.03839791},"labels":[],"label_agreement":null},{"id":"W2915189456","doi":"10.1155/2019/8367214","title":"Software Development Effort Estimation Using Regression Fuzzy Models","year":2019,"lang":"en","type":"article","venue":"Computational Intelligence and Neuroscience","topic":"Software Engineering Research","field":"Computer Science","cited_by":184,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Western University","funders":"American University of Sharjah; University of Sharjah; Applied Science Private University","keywords":"Computer science; Fuzzy logic; Machine learning; Data mining; Outlier; Artificial intelligence; Benchmarking; Heteroscedasticity; Software","score_opus":0.07640391814971134,"score_gpt":0.32333245734537525,"score_spread":0.2469285391956639,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2915189456","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1988173,0.0005270047,0.797575,0.00019759302,0.000028007285,0.00010889834,0.00020817586,0.00055048673,0.0019875525],"genre_scores_gemma":[0.889613,0.00030730054,0.108679816,0.000031866384,0.0000160296,0.00012382788,0.00019952956,0.000025482583,0.0010032111],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99830127,0.0007483289,0.0001238972,0.00035594674,0.00038352347,0.000086996675],"domain_scores_gemma":[0.99371135,0.004515523,0.0006119628,0.00020782184,0.00090195984,0.000051396997],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035162752,0.00085780094,0.000871807,0.0017652233,0.00030898338,0.0011951568,0.001064633,0.0007451519,0.00080337364],"category_scores_gemma":[0.012892273,0.00042349152,0.0009952717,0.0012920802,0.00028224228,0.001223672,0.00050777034,0.00091134856,0.00021546402],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009811996,0.00007231863,0.0036454338,0.00007763024,0.00007965329,0.000037868795,0.0000821356,0.9447218,0.0011913615,0.0025440895,0.00023480025,0.047214735],"study_design_scores_gemma":[0.0000036375911,0.000023802446,0.00045476964,0.0000073678752,0.0000094190655,0.0000039721017,0.000007723332,0.9982204,0.0004078685,0.0007690378,0.00008644554,0.0000055954692],"about_ca_topic_score_codex":0.012822916,"about_ca_topic_score_gemma":0.0077343816,"teacher_disagreement_score":0.012822916,"about_ca_system_score_codex":0.001047659,"about_ca_system_score_gemma":0.00089699443,"threshold_uncertainty_score":0.025496602},"labels":[],"label_agreement":null},{"id":"W2915353454","doi":"10.1145/1095430.1095433","title":"Report on MSR 2005","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Session (web analytics); Plan (archaeology); Software; Computer science; Field (mathematics); Software review; Quality (philosophy); Software engineering; World Wide Web; Software development; Software construction; History; Operating system","score_opus":0.01767719627848747,"score_gpt":0.2610800800871416,"score_spread":0.24340288380865416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2915353454","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008725662,0.011563009,0.005962433,0.053861957,0.098889686,0.0028998598,0.084097356,0.01059707,0.7234029],"genre_scores_gemma":[0.0108734695,0.0029508572,0.0029821594,0.008169971,0.0066777864,0.00078759086,0.04302276,0.0016748455,0.9228605],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99290013,0.0008149457,0.00051370007,0.0008376232,0.0041801394,0.00075354066],"domain_scores_gemma":[0.9866484,0.0008354088,0.00053164596,0.0015103593,0.0076428265,0.0028314108],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.008224122,0.0015274377,0.0012447343,0.0036318018,0.0024779732,0.008887058,0.0027854564,0.0038747503,0.3737671],"category_scores_gemma":[0.012924888,0.00048234625,0.0015085787,0.0024333596,0.00037970298,0.0032610372,0.0040292754,0.0029285485,0.36794367],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016358228,0.00007434123,0.00029373463,0.00011987743,0.000008158329,0.00008565734,0.00004922759,0.00005258818,0.0004140531,0.0012313485,0.9673519,0.030155506],"study_design_scores_gemma":[0.000015384428,0.000066614964,0.0007259273,0.000057934398,0.0000042937413,0.000027255395,0.00005882081,0.000037908205,0.00030567375,0.00019610376,0.9984969,0.000007195939],"about_ca_topic_score_codex":0.006105847,"about_ca_topic_score_gemma":0.008007035,"teacher_disagreement_score":0.3737671,"about_ca_system_score_codex":0.003690026,"about_ca_system_score_gemma":0.006208625,"threshold_uncertainty_score":0.8932452},"labels":[],"label_agreement":null},{"id":"W2915705208","doi":"","title":"Proceedings of the 7th International Workshop on Software Clones","year":2013,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; clone (Java method); Restructuring; Theme (computing); Software engineering; Software; Software development; Software maintenance; Computer science; Engineering management; Engineering; World Wide Web; Programming language; Business; Biology","score_opus":0.026868047668836027,"score_gpt":0.26482850738406744,"score_spread":0.23796045971523141,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2915705208","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021481099,0.11227279,0.35471705,0.063018315,0.07828645,0.0014925181,0.0026866712,0.0084186215,0.35762653],"genre_scores_gemma":[0.10826565,0.084767945,0.24877918,0.012662595,0.020636413,0.0019659644,0.01847693,0.0062712673,0.49817404],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99320376,0.0023275125,0.00060267356,0.0010844846,0.0021226679,0.0006589523],"domain_scores_gemma":[0.9905344,0.0029918747,0.00031641463,0.0018712821,0.0030377745,0.0012482628],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007756061,0.001638693,0.0015126402,0.0029524039,0.0018301489,0.0082772635,0.0034313288,0.003778966,0.058921285],"category_scores_gemma":[0.015325131,0.0010151791,0.0019200806,0.0025618759,0.002046461,0.009528397,0.0067577963,0.0049302713,0.022205552],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029460836,0.00028356913,0.0008487281,0.00082083663,0.0000753322,0.0005812798,0.001960094,0.0019530513,0.0039360123,0.04502073,0.53647417,0.4077517],"study_design_scores_gemma":[0.000031630032,0.00006802233,0.00054131553,0.00057409453,0.000036405425,0.00046388022,0.00050573755,0.0018513704,0.0010702563,0.014939582,0.979889,0.000028638866],"about_ca_topic_score_codex":0.0029700242,"about_ca_topic_score_gemma":0.0038264643,"teacher_disagreement_score":0.058921285,"about_ca_system_score_codex":0.0022645397,"about_ca_system_score_gemma":0.0034628953,"threshold_uncertainty_score":0.19711131},"labels":[],"label_agreement":null},{"id":"W2916792555","doi":"","title":"Proceedings of the 5th Workshop on Evaluation and Usability of Programming Languages and Tools","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Calgary","funders":"","keywords":"Usability; Computer science; Presentation (obstetrics); Software engineering; World Wide Web; Programming language; Human–computer interaction","score_opus":0.03449683749549182,"score_gpt":0.3356622200430833,"score_spread":0.3011653825475915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2916792555","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034164123,0.055383306,0.50038147,0.06214755,0.025904952,0.008026254,0.0044911085,0.01012686,0.29937434],"genre_scores_gemma":[0.1721275,0.029106038,0.49561083,0.012830544,0.005751823,0.012508795,0.016544718,0.010541912,0.24497782],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9360405,0.03997635,0.004655369,0.004242195,0.013091303,0.0019942464],"domain_scores_gemma":[0.91792095,0.03867668,0.0021849968,0.010024554,0.026569761,0.004623085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06054405,0.00223175,0.0021433074,0.0035030092,0.0026602005,0.016516158,0.0041883094,0.0038290443,0.048189364],"category_scores_gemma":[0.081302024,0.0015604709,0.0022817291,0.002159249,0.0039709946,0.012805712,0.009638262,0.006465306,0.016291738],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073721254,0.00078629213,0.0025679313,0.0022741498,0.000228833,0.0003207651,0.009289027,0.0011300902,0.0052858046,0.028856393,0.3903947,0.5581289],"study_design_scores_gemma":[0.00018567318,0.00055526505,0.004762224,0.0043460373,0.00016444319,0.00041929356,0.003244629,0.0023851902,0.004048973,0.025323091,0.95441574,0.00014944379],"about_ca_topic_score_codex":0.004933749,"about_ca_topic_score_gemma":0.0048453086,"teacher_disagreement_score":0.06054405,"about_ca_system_score_codex":0.0038607165,"about_ca_system_score_gemma":0.007132965,"threshold_uncertainty_score":0.32019138},"labels":[],"label_agreement":null},{"id":"W2917655874","doi":"10.48550/arxiv.1902.07093","title":"Analysis and Detection of Information Types of Open Source Software Issue Discussions","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Task (project management); Sentence; Software; Data science; World Wide Web; Artificial intelligence; Engineering","score_opus":0.030600418984793022,"score_gpt":0.20686761001188128,"score_spread":0.17626719102708827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2917655874","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9488331,0.00078788644,0.038945217,0.00070097926,0.00012049862,0.00058903696,0.0049218186,0.0006624436,0.004439057],"genre_scores_gemma":[0.9114016,0.0005096339,0.07530266,0.00021009456,0.00016730273,0.0010026512,0.00855052,0.00023382908,0.0026216863],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99186033,0.003942327,0.0008491612,0.0011677477,0.0017928659,0.0003875483],"domain_scores_gemma":[0.9205874,0.057101082,0.009120053,0.0023155515,0.009777355,0.0010984818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005925511,0.00051049946,0.00045480643,0.00994788,0.0014059091,0.0018438945,0.0007060001,0.0008682784,0.0011318959],"category_scores_gemma":[0.04432345,0.00036182316,0.0004462422,0.004790495,0.00094021345,0.0031169592,0.002104277,0.0010145593,0.00054164673],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017407751,0.00059086754,0.2916124,0.0051448215,0.00017738625,0.0024944346,0.174912,0.0021392298,0.12397075,0.011903866,0.015309029,0.3700044],"study_design_scores_gemma":[0.00011574958,0.00072295143,0.5911618,0.0013755234,0.00034603465,0.0031929417,0.089091174,0.10364392,0.084427595,0.016970446,0.10852884,0.00042296937],"about_ca_topic_score_codex":0.002512692,"about_ca_topic_score_gemma":0.0037362399,"teacher_disagreement_score":0.00994788,"about_ca_system_score_codex":0.0013345483,"about_ca_system_score_gemma":0.0014129364,"threshold_uncertainty_score":0.03133744},"labels":[],"label_agreement":null},{"id":"W2918769567","doi":"","title":"Just-in-Time Detection of Protection-Impacting Changes on Wordpress and Mediawiki","year":2018,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Humanities; Political science; Computer science; Art","score_opus":0.016422666058007896,"score_gpt":0.2481954944721201,"score_spread":0.2317728284141122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2918769567","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9330929,0.00092372403,0.037480846,0.0002895311,0.00014752419,0.00024933778,0.0034998087,0.021754751,0.0025615809],"genre_scores_gemma":[0.94593155,0.00027485978,0.042172898,0.00011222056,0.000025759258,0.00028777518,0.004493821,0.0015713287,0.0051298607],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99564064,0.00044271766,0.00038796797,0.000970147,0.0021754326,0.00038301054],"domain_scores_gemma":[0.98414814,0.0065981927,0.0026357404,0.0027636723,0.003299543,0.00055478426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015715095,0.00080527365,0.00054392725,0.0024116968,0.0006101418,0.0011068512,0.0008646758,0.0007653634,0.0014074417],"category_scores_gemma":[0.01588479,0.0006409416,0.00048728567,0.0020086132,0.0008846599,0.002688985,0.0010394269,0.0013264765,0.00070105644],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036649287,0.00061989663,0.2128449,0.0024324316,0.00041170075,0.003958844,0.00782049,0.012404468,0.25473335,0.002866351,0.014464314,0.48377827],"study_design_scores_gemma":[0.00013412353,0.0015183381,0.37827915,0.00019387224,0.0002590914,0.0023873525,0.0020383412,0.21222045,0.36282057,0.0034208547,0.036437172,0.000290713],"about_ca_topic_score_codex":0.008115514,"about_ca_topic_score_gemma":0.011594373,"teacher_disagreement_score":0.008115514,"about_ca_system_score_codex":0.00092192885,"about_ca_system_score_gemma":0.001333422,"threshold_uncertainty_score":0.016136527},"labels":[],"label_agreement":null},{"id":"W2919248872","doi":"10.1016/j.jss.2019.02.056","title":"A survey of self-admitted technical debt","year":2019,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Japan Society for the Promotion of Science","keywords":"Technical debt; Replicate; Computer science; Comprehension; Data science; Set (abstract data type); Work (physics); Debt; Empirical research; Code (set theory); Implementation; Source code; Risk analysis (engineering); Data mining; Software engineering; Software; Engineering; Business; Finance; Software development; Statistics","score_opus":0.016736680814053474,"score_gpt":0.2568966118747613,"score_spread":0.24015993106070785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2919248872","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88288814,0.043265365,0.0054794303,0.009049226,0.00026612845,0.00006823744,0.0033131044,0.00014792544,0.055522323],"genre_scores_gemma":[0.95009315,0.035474986,0.002271914,0.001176568,0.0003288526,0.000040431507,0.0032299648,0.000066679684,0.0073174834],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9974146,0.0006353889,0.00038575198,0.00022439683,0.000972801,0.00036697777],"domain_scores_gemma":[0.97368026,0.010151588,0.007075565,0.0013979794,0.005383129,0.0023113482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036121586,0.00019955933,0.00034282118,0.006351188,0.0010464839,0.0026503042,0.0008884277,0.00082380674,0.006132053],"category_scores_gemma":[0.027318064,0.0002197928,0.00029617446,0.010521806,0.0010091162,0.004128926,0.0015292467,0.0011309049,0.0009889845],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002430654,0.00026520572,0.68525046,0.0004803879,0.000057718502,0.0006730447,0.003816984,0.00042380404,0.0003894298,0.033206735,0.017819643,0.25737357],"study_design_scores_gemma":[0.000029077573,0.000383111,0.6809867,0.0014691547,0.00008542772,0.0073439563,0.012868165,0.0031596415,0.0012249631,0.029943356,0.2624222,0.00008425991],"about_ca_topic_score_codex":0.00408623,"about_ca_topic_score_gemma":0.0037189813,"teacher_disagreement_score":0.006351188,"about_ca_system_score_codex":0.0012835977,"about_ca_system_score_gemma":0.0018599833,"threshold_uncertainty_score":0.020513833},"labels":[],"label_agreement":null},{"id":"W2919910741","doi":"10.1016/j.jss.2019.110486","title":"A machine-learning based ensemble method for anti-patterns detection","year":2019,"lang":"en","type":"preprint","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Computer science; Ensemble learning; Machine learning; False positive paradox; Maintainability; Artificial intelligence; Aggregate (composite); Feature (linguistics); Source code; Program comprehension; Java; Class (philosophy); Data mining; Code (set theory); Pattern recognition (psychology); Software; Set (abstract data type); Software system","score_opus":0.024625563804385794,"score_gpt":0.2921812808401903,"score_spread":0.2675557170358045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2919910741","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016476963,0.0003626863,0.98095393,0.0000872688,0.00014504277,0.00004674433,0.00012683471,0.001122549,0.00067792466],"genre_scores_gemma":[0.26082355,0.00042150053,0.73056346,0.00020142127,0.0002622903,0.0001579054,0.0010670086,0.0002856588,0.006217121],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99825305,0.00034968575,0.00012776349,0.0004198403,0.0006643576,0.00018521368],"domain_scores_gemma":[0.99672157,0.0010636364,0.00016095834,0.0006014239,0.0013242327,0.000128165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022037902,0.0010467682,0.0019507117,0.00219113,0.00088381185,0.0010756114,0.0019660299,0.0016518161,0.0024283703],"category_scores_gemma":[0.0041415347,0.00046619878,0.0012857468,0.0021530625,0.00034327628,0.0016136534,0.0013814364,0.0015570564,0.0014102779],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026920487,0.00027947532,0.0033502432,0.00007787506,0.00030141638,0.000118444885,0.000074081625,0.049495086,0.02347085,0.0022493864,0.0051141097,0.9151998],"study_design_scores_gemma":[0.0000085722195,0.00007458823,0.0011255518,0.000007432362,0.000060407478,0.0001241569,0.00001461362,0.98979455,0.00586114,0.0016122854,0.001300469,0.000016113232],"about_ca_topic_score_codex":0.0023329728,"about_ca_topic_score_gemma":0.0045621987,"teacher_disagreement_score":0.0024283703,"about_ca_system_score_codex":0.00029036563,"about_ca_system_score_gemma":0.0007927751,"threshold_uncertainty_score":0.011654913},"labels":[],"label_agreement":null},{"id":"W2920526824","doi":"10.1007/s10664-019-09695-9","title":"An empirical study of the long duration of continuous integration builds","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":82,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Duration (music); Computer science; Set (abstract data type); Empirical research; Software development; Cache; Software; Software engineering; Process management; Engineering; Statistics; Operating system; Programming language","score_opus":0.016944779324193943,"score_gpt":0.2973705753052279,"score_spread":0.280425795981034,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2920526824","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9938267,0.00044293577,0.0009305709,0.00027331413,0.000010512202,0.000014942018,0.00012132221,0.000018588977,0.004360995],"genre_scores_gemma":[0.99891543,0.00008873181,0.00024330943,0.000029043515,0.000013790086,0.000009839599,0.000110230394,0.000007832455,0.0005818794],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9978466,0.0007039348,0.00015201532,0.00027117654,0.00077824836,0.00024791397],"domain_scores_gemma":[0.82757074,0.1255152,0.027043065,0.0063669286,0.00743384,0.0060701943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052403007,0.00019149082,0.00026113354,0.0013831147,0.0011793369,0.0025686766,0.0010296764,0.0010469686,0.0071107224],"category_scores_gemma":[0.07605561,0.00027619922,0.00020758006,0.0023573313,0.0012128501,0.0034088078,0.0018285424,0.0020077599,0.00073262275],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014391749,0.0016558745,0.8718597,0.00026344808,0.00013646408,0.0007902127,0.014317352,0.0016224755,0.002770205,0.01442274,0.0014666948,0.089255475],"study_design_scores_gemma":[0.000091879956,0.0015740215,0.9559712,0.00013733881,0.00015538324,0.00075946125,0.016751532,0.0053416993,0.0015385112,0.009628538,0.007986527,0.00006396296],"about_ca_topic_score_codex":0.0030757294,"about_ca_topic_score_gemma":0.0048115104,"teacher_disagreement_score":0.0071107224,"about_ca_system_score_codex":0.001044201,"about_ca_system_score_gemma":0.0012643309,"threshold_uncertainty_score":0.027713716},"labels":[],"label_agreement":null},{"id":"W2920574268","doi":"","title":"Predictive analytics in healthcare epileptic seizure recognition","year":2018,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Epilepsy; Epileptic seizure; Electroencephalography; Computer science; Artificial intelligence; Binary classification; Machine learning; Random forest; Predictive analytics; Pattern recognition (psychology); Support vector machine; Psychology; Psychiatry","score_opus":0.019426985438490883,"score_gpt":0.25326440946288054,"score_spread":0.23383742402438965,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2920574268","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12522532,0.027469555,0.7846967,0.017775955,0.00091020623,0.0004804507,0.007781409,0.00785276,0.027807694],"genre_scores_gemma":[0.90335655,0.007326362,0.080758885,0.00064332684,0.00044520132,0.00015525788,0.0040513165,0.000101464655,0.0031616339],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99901426,0.00030957148,0.00008495422,0.00017123003,0.00034903453,0.00007091019],"domain_scores_gemma":[0.9963038,0.0024408975,0.0003626669,0.0002365798,0.0005824399,0.00007361417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017314911,0.000734331,0.0007335952,0.0023889472,0.00027373654,0.0018542901,0.00068759033,0.00065639254,0.002849342],"category_scores_gemma":[0.0067865434,0.00016799006,0.0004423283,0.0027142884,0.0004028353,0.00128219,0.00076766755,0.0010113616,0.00093628775],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021050818,0.0002500507,0.024303561,0.0006413727,0.00016458592,0.00042512995,0.00018879407,0.25899118,0.0021189242,0.028782966,0.028383078,0.6555398],"study_design_scores_gemma":[0.00001135151,0.00006252978,0.0054928744,0.00019073699,0.0000312661,0.00015016017,0.00017704602,0.91758704,0.0017669594,0.06421257,0.010291334,0.000026150423],"about_ca_topic_score_codex":0.006680838,"about_ca_topic_score_gemma":0.0035057964,"teacher_disagreement_score":0.006680838,"about_ca_system_score_codex":0.00089472684,"about_ca_system_score_gemma":0.0007568723,"threshold_uncertainty_score":0.013283908},"labels":[],"label_agreement":null},{"id":"W2921039133","doi":"10.1109/saner.2019.8667999","title":"Is Self-Admitted Technical Debt a Good Indicator of Architectural Divergences?","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Technical debt; Architectural pattern; Computer science; TRACE (psycholinguistics); Implementation; Source code; Leverage (statistics); Documentation; Software engineering; Code (set theory); Software; Set (abstract data type); Data science; Software design; Programming language; Software development; Artificial intelligence","score_opus":0.00922187233238762,"score_gpt":0.2521252984220785,"score_spread":0.24290342608969087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2921039133","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9654537,0.00056792557,0.025968244,0.000891924,0.00006205997,0.000102578306,0.0010082953,0.00068040594,0.00526497],"genre_scores_gemma":[0.9768067,0.00017517962,0.01983669,0.00023871985,0.00005429674,0.000073288895,0.0013849307,0.00025583117,0.001174395],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98212206,0.004033243,0.002838122,0.002342083,0.007854866,0.0008096508],"domain_scores_gemma":[0.6414499,0.14091411,0.118011475,0.04000806,0.05355715,0.0060591507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011530139,0.0005492977,0.0005255673,0.0076061715,0.0013850828,0.0025217673,0.0011499544,0.0015573025,0.0016631307],"category_scores_gemma":[0.16370949,0.00061092206,0.0004190189,0.0057826857,0.0019912035,0.005893594,0.0038334373,0.0018527762,0.0006163791],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022989261,0.0001365354,0.89737284,0.0005193037,0.00007241809,0.00078711164,0.012478626,0.0010023311,0.0076786666,0.0034369265,0.0015990497,0.07468627],"study_design_scores_gemma":[0.000027245942,0.0003786092,0.9239698,0.0006140536,0.000103156795,0.0027149234,0.014720817,0.0145180235,0.007265918,0.014754675,0.020751985,0.00018085583],"about_ca_topic_score_codex":0.00296304,"about_ca_topic_score_gemma":0.0035656248,"teacher_disagreement_score":0.011530139,"about_ca_system_score_codex":0.0013303396,"about_ca_system_score_gemma":0.0013914693,"threshold_uncertainty_score":0.060977995},"labels":[],"label_agreement":null},{"id":"W2921166721","doi":"10.1109/saner.2019.8668029","title":"Reuse (or Lack Thereof) in Travis CI Specifications: An Empirical Study of CI Phases and Commands","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Reuse; Java; Computer science; Lift (data mining); Software engineering; Association rule learning; Sample (material); Empirical research; World Wide Web; Programming language; Data mining; Engineering","score_opus":0.16362606862881607,"score_gpt":0.39279492327660576,"score_spread":0.2291688546477897,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2921166721","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99795544,0.000111785834,0.0010428313,0.00006230873,0.0000025284905,0.000036610803,0.00015304673,0.000039208982,0.0005962247],"genre_scores_gemma":[0.996639,0.0001230342,0.0021089318,0.00003552995,0.000004942746,0.00008028972,0.0005358499,0.000054881748,0.0004174869],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9781385,0.006251208,0.0033382282,0.0023650394,0.008702863,0.0012040323],"domain_scores_gemma":[0.66346157,0.20791058,0.07523735,0.016130703,0.032632787,0.004626977],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.020031154,0.00039239603,0.0004795817,0.0053560766,0.0010382072,0.0030965158,0.0014475542,0.0008583962,0.0010535662],"category_scores_gemma":[0.13602886,0.0005650497,0.00045288337,0.006233548,0.0023352837,0.0043516112,0.0028626428,0.0015283005,0.00039345853],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002136371,0.00028160246,0.92957103,0.00028182165,0.000069372574,0.0004237505,0.0262645,0.00047773946,0.0018980869,0.0006736159,0.0006140459,0.039230842],"study_design_scores_gemma":[0.0000224104,0.0006469895,0.949554,0.0002443695,0.00006370998,0.0012540754,0.034349527,0.0053109527,0.0027819215,0.00060667295,0.0050847726,0.00008053704],"about_ca_topic_score_codex":0.006728972,"about_ca_topic_score_gemma":0.009162788,"teacher_disagreement_score":0.97996885,"about_ca_system_score_codex":0.0018328043,"about_ca_system_score_gemma":0.0024567891,"threshold_uncertainty_score":0.10593611},"labels":[],"label_agreement":null},{"id":"W2921412691","doi":"10.1109/saner.2019.8667993","title":"A Comparative Study of Software Bugs in Micro-clones and Regular Code Clones","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Software bug; Programming language; Code (set theory); Operating system; Software","score_opus":0.022690538973877135,"score_gpt":0.2877814240137133,"score_spread":0.2650908850398362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2921412691","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9968736,0.00026456884,0.0024743003,0.00001370869,0.0000038394846,0.000019588984,0.00006354195,0.000079318816,0.00020743937],"genre_scores_gemma":[0.99545705,0.00010094382,0.003976636,0.0000114996465,0.0000054816114,0.000018551322,0.00020892003,0.000020581161,0.00020039173],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9965887,0.0005835893,0.00034775963,0.0008034063,0.0014597324,0.00021681699],"domain_scores_gemma":[0.9496486,0.030843643,0.009920206,0.0031624488,0.005149447,0.0012756141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016440668,0.00030261313,0.00040296474,0.0035565828,0.0003643783,0.00055185286,0.00042882774,0.0004945884,0.00057841116],"category_scores_gemma":[0.021703787,0.0001859965,0.0003116567,0.0017455268,0.00073293695,0.0012965173,0.0007405543,0.00032242175,0.00008052604],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013024777,0.000468114,0.75412804,0.0008436679,0.00031623745,0.0025261228,0.006032554,0.0043803575,0.067604475,0.0010834006,0.0005637332,0.16075076],"study_design_scores_gemma":[0.000050108338,0.0028162727,0.94155157,0.000089875335,0.00025498468,0.0060436837,0.0026432818,0.020198353,0.02343509,0.0010403756,0.0018038538,0.000072633076],"about_ca_topic_score_codex":0.0009423506,"about_ca_topic_score_gemma":0.0015420369,"teacher_disagreement_score":0.0035565828,"about_ca_system_score_codex":0.00033921286,"about_ca_system_score_gemma":0.0002832564,"threshold_uncertainty_score":0.008694708},"labels":[],"label_agreement":null},{"id":"W2921727127","doi":"10.1109/saner.2019.8668042","title":"Detecting Feature-Interaction Symptoms in Automotive Software using Lightweight Analysis","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Toolchain; Codebase; Computer science; Feature (linguistics); Source code; Set (abstract data type); Software; Code (set theory); Automotive industry; Static program analysis; Data mining; Programming language; Software engineering; Software development; Engineering","score_opus":0.012253024479678682,"score_gpt":0.2760089249537191,"score_spread":0.2637559004740404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2921727127","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57638013,0.00049185177,0.40124112,0.0005596975,0.000027185426,0.00032138993,0.0019508478,0.017276648,0.0017510734],"genre_scores_gemma":[0.76942235,0.00021483988,0.2254144,0.00013364415,0.000014978566,0.0001546168,0.003319746,0.0005837217,0.0007417276],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958407,0.0006683281,0.0004320358,0.0008032486,0.0020451285,0.0002105743],"domain_scores_gemma":[0.98151416,0.010304784,0.0034159045,0.002152597,0.002396518,0.0002161089],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020334318,0.00080535846,0.000457969,0.005921844,0.00058286916,0.0013453779,0.0013805459,0.0007429729,0.00069455296],"category_scores_gemma":[0.016235055,0.0005259329,0.001160545,0.002813283,0.0008716343,0.0021172422,0.0014236226,0.00082616747,0.00029848446],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043292774,0.00058060617,0.3136415,0.0011237165,0.00035797412,0.0043808217,0.0071088155,0.05925978,0.1100869,0.010916798,0.0063991025,0.4857112],"study_design_scores_gemma":[0.000070489914,0.000515537,0.102927975,0.0002763794,0.00032080742,0.00219703,0.0015906963,0.75864404,0.10293118,0.015079255,0.015289931,0.00015667966],"about_ca_topic_score_codex":0.007597243,"about_ca_topic_score_gemma":0.010396666,"teacher_disagreement_score":0.007597243,"about_ca_system_score_codex":0.0010274481,"about_ca_system_score_gemma":0.0016548502,"threshold_uncertainty_score":0.015106022},"labels":[],"label_agreement":null},{"id":"W2921931467","doi":"10.1016/j.scico.2019.03.001","title":"Measuring and analyzing code authorship in 1 + 118 open source projects","year":2019,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Fundação de Amparo à Pesquisa do Estado de Minas Gerais; Conselho Nacional de Desenvolvimento Científico e Tecnológico; Coordenação de Aperfeiçoamento de Pessoal de Nível Superior","keywords":"Computer science; Linux kernel; Source code; Relation (database); Key (lock); World Wide Web; Code (set theory); Open source; Replicate; Constant (computer programming); Metric (unit); Open source software; Operating system; Software; Information retrieval; Data science; Programming language; Database; Statistics","score_opus":0.05623410123772979,"score_gpt":0.3030603189937094,"score_spread":0.2468262177559796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2921931467","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9964954,0.00013020367,0.0013816751,0.00007554125,0.000019025989,0.000030765063,0.00033603903,0.00013463614,0.0013967658],"genre_scores_gemma":[0.99038965,0.0000911824,0.0061338414,0.000021758,0.000025697213,0.00006699035,0.0011889202,0.000087333574,0.001994513],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98801047,0.003102364,0.0015095768,0.0019200959,0.004844158,0.0006132584],"domain_scores_gemma":[0.8395121,0.083979785,0.031018572,0.010940313,0.028333504,0.0062157344],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.006131408,0.00035815433,0.0003165389,0.0070846053,0.001339801,0.0014655392,0.0007120893,0.00084309734,0.0011345455],"category_scores_gemma":[0.08011086,0.00029808548,0.0002843162,0.0044508865,0.0008219175,0.0019428576,0.0025302651,0.0006961812,0.00052158255],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041727116,0.0002835737,0.8166114,0.00030203254,0.000082057704,0.000292636,0.0060080322,0.0014759551,0.006520628,0.0014893269,0.0019562966,0.1645608],"study_design_scores_gemma":[0.00007380224,0.0007398699,0.9376845,0.0001596703,0.00011155323,0.0012413362,0.005369612,0.021475442,0.0111484,0.0037535965,0.018157514,0.00008468202],"about_ca_topic_score_codex":0.0014645647,"about_ca_topic_score_gemma":0.0041332007,"teacher_disagreement_score":0.9938686,"about_ca_system_score_codex":0.00079549674,"about_ca_system_score_gemma":0.0013542314,"threshold_uncertainty_score":0.032426357},"labels":[],"label_agreement":null},{"id":"W2922277646","doi":"10.1109/icpc.2019.00054","title":"Recommending Comprehensive Solutions for Programming Tasks by Mining Crowd Knowledge","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Canada First Research Excellence Fund","keywords":"Computer science; Code (set theory); Relevance (law); Task (project management); Information retrieval; Quality (philosophy); Programming language; Artificial intelligence","score_opus":0.08518891435909169,"score_gpt":0.3441415631625433,"score_spread":0.2589526488034516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2922277646","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31839278,0.0062130853,0.59301883,0.004182304,0.0005859944,0.0032497412,0.019268617,0.0257577,0.02933084],"genre_scores_gemma":[0.4076558,0.0016175758,0.5395454,0.0012837565,0.0002841282,0.0015892623,0.037262842,0.0009949746,0.009766173],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99624985,0.0008101781,0.00027350595,0.001095065,0.0013584924,0.00021293414],"domain_scores_gemma":[0.9942287,0.0035273063,0.000479369,0.0004082092,0.0010174856,0.0003388963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022369938,0.0034583502,0.001646364,0.013042434,0.0013289985,0.0021607184,0.0026504013,0.0032517556,0.005838195],"category_scores_gemma":[0.0146736335,0.0008448126,0.0018864988,0.0044673686,0.0007790613,0.004387458,0.003054266,0.001587993,0.0031080004],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013662761,0.0018082042,0.021473058,0.0048849224,0.00045596558,0.0029383432,0.004122558,0.037288256,0.020430878,0.008821409,0.11750019,0.77891004],"study_design_scores_gemma":[0.00071128586,0.0009761838,0.015772112,0.00096322404,0.00072127156,0.0016271342,0.009011412,0.74772066,0.027099937,0.055058636,0.14000045,0.00033769663],"about_ca_topic_score_codex":0.0077539743,"about_ca_topic_score_gemma":0.017027088,"teacher_disagreement_score":0.013042434,"about_ca_system_score_codex":0.0013731073,"about_ca_system_score_gemma":0.0028111995,"threshold_uncertainty_score":0.019530654},"labels":[],"label_agreement":null},{"id":"W2922601124","doi":"10.1016/j.visinf.2019.03.003","title":"Clone-World: A visual analytic system for large scale software clones","year":2019,"lang":"en","type":"article","venue":"Visual Informatics","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Canada First Research Excellence Fund","keywords":"Code refactoring; clone (Java method); Computer science; Software; Software maintenance; Software system; Software development; Software visualization; Software engineering; Visual analytics; Data science; Visualization; Data mining; Software construction; Programming language; Biology","score_opus":0.01327765827858923,"score_gpt":0.30516560612500565,"score_spread":0.2918879478464164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2922601124","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010807996,0.00030113038,0.85689837,0.00045159974,0.00008029841,0.0004528856,0.003337849,0.12507945,0.0025904258],"genre_scores_gemma":[0.13864963,0.0006187225,0.84037733,0.00033174225,0.00011042421,0.0011912298,0.0059429663,0.010422405,0.0023555984],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990619,0.00028223844,0.000101550904,0.00018374033,0.0003145425,0.000056066445],"domain_scores_gemma":[0.99276745,0.004142661,0.0006014841,0.0010666008,0.000984771,0.0004370351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003002451,0.001616032,0.00093136565,0.0061980784,0.00072266214,0.0033111589,0.0023426965,0.0012848703,0.012030016],"category_scores_gemma":[0.014654983,0.0008257606,0.0011241013,0.0023852761,0.00066497596,0.005167496,0.004871038,0.0016228156,0.0019698322],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018700778,0.0004433489,0.0126034245,0.0016519487,0.00042661154,0.0013124717,0.007318823,0.030797781,0.02787941,0.053104393,0.11899451,0.7435972],"study_design_scores_gemma":[0.00040664538,0.00035258036,0.0074525042,0.0004918645,0.00018535205,0.0010771152,0.001534409,0.70331925,0.026440995,0.104591265,0.15381253,0.00033555526],"about_ca_topic_score_codex":0.002342116,"about_ca_topic_score_gemma":0.002297232,"teacher_disagreement_score":0.012030016,"about_ca_system_score_codex":0.00081656536,"about_ca_system_score_gemma":0.00082263944,"threshold_uncertainty_score":0.0402444},"labels":[],"label_agreement":null},{"id":"W2924144435","doi":"10.1007/s11219-019-09442-9","title":"A large-scale empirical study of code smells in JavaScript projects","year":2019,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code smell; JavaScript; Computer science; Scripting language; Software quality; World Wide Web; Programming language; Software development; Software","score_opus":0.05973784110211131,"score_gpt":0.36947289628882407,"score_spread":0.30973505518671274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2924144435","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9994722,0.000027586713,0.00016534975,0.000036361092,0.0000013186526,0.000008397681,0.00006851893,0.0000072897024,0.00021301907],"genre_scores_gemma":[0.9991885,0.000029990259,0.00032703968,0.000028249246,0.0000042717356,0.000018015266,0.00021293928,0.000009261621,0.00018174587],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9962049,0.0013733964,0.00033097397,0.00055049453,0.0012299521,0.00031023176],"domain_scores_gemma":[0.86367196,0.082127884,0.030408308,0.006348079,0.01174616,0.005697606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053270487,0.0002494265,0.00026097224,0.0018875576,0.00078747346,0.0011913793,0.0006998732,0.0009205723,0.0009718423],"category_scores_gemma":[0.049172115,0.00035637847,0.0003603538,0.0023448062,0.0010650133,0.0027699368,0.0014195397,0.0015196716,0.0004385737],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013144236,0.001026372,0.9827628,0.000059534847,0.00006063784,0.00015842564,0.0046019903,0.00027111225,0.0010946674,0.00015127231,0.000451773,0.009230005],"study_design_scores_gemma":[0.000011728349,0.00026043106,0.99337935,0.00002631761,0.000016774597,0.00011310425,0.0037391833,0.0016226796,0.0003405706,0.00009471579,0.00038008785,0.000014864997],"about_ca_topic_score_codex":0.0063783554,"about_ca_topic_score_gemma":0.013891428,"teacher_disagreement_score":0.0063783554,"about_ca_system_score_codex":0.0008843442,"about_ca_system_score_gemma":0.0011715435,"threshold_uncertainty_score":0.028172433},"labels":[],"label_agreement":null},{"id":"W2929132723","doi":"10.24251/hicss.2019.892","title":"Examining User-Developer Feedback Loops in the iOS App Store","year":2019,"lang":"en","type":"article","venue":"Proceedings of the ... Annual Hawaii International Conference on System Sciences/Proceedings of the Annual Hawaii International Conference on System Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Computer science; Login; Focus (optics); Feedback loop; World Wide Web; App store; Feature (linguistics); Human–computer interaction; Information retrieval; Computer security","score_opus":0.05509498982379558,"score_gpt":0.29437606310841397,"score_spread":0.2392810732846184,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2929132723","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9933009,0.00046800746,0.003958942,0.00016458881,0.000017741519,0.00011971693,0.00075489155,0.00032386582,0.0008912883],"genre_scores_gemma":[0.9873148,0.00023636881,0.007991394,0.00012613214,0.000027458287,0.00022955029,0.0028142568,0.00013725925,0.001122657],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9906544,0.0035631869,0.0009313363,0.0016628748,0.002811389,0.00037679213],"domain_scores_gemma":[0.79420155,0.16271354,0.023960145,0.004138456,0.013452917,0.0015333669],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008396168,0.00053174427,0.0005193316,0.0064119487,0.0010683301,0.0017440385,0.00074363913,0.000928177,0.0005767302],"category_scores_gemma":[0.08803979,0.0005360152,0.00036323827,0.002975956,0.00094288634,0.003061338,0.0015233627,0.0010300777,0.00041150546],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078942033,0.0007638987,0.7609323,0.0014429644,0.00021262334,0.0018990703,0.06392794,0.0016834534,0.016047234,0.0008425206,0.009055523,0.14240308],"study_design_scores_gemma":[0.000054587377,0.00048594357,0.91502154,0.00034173386,0.00012335643,0.0017781113,0.013525615,0.04070674,0.009456343,0.0012106613,0.01713058,0.00016487812],"about_ca_topic_score_codex":0.009384688,"about_ca_topic_score_gemma":0.022608459,"teacher_disagreement_score":0.009384688,"about_ca_system_score_codex":0.0013007673,"about_ca_system_score_gemma":0.0013206286,"threshold_uncertainty_score":0.044403672},"labels":[],"label_agreement":null},{"id":"W2932400032","doi":"10.1145/3302509.3313318","title":"Feature characterization for CPS software reuse","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Reuse; Feature (linguistics); Abstraction; Cyber-physical system; Software engineering; Software; Software system; Programming language; Engineering","score_opus":0.010900022491098156,"score_gpt":0.24395257573780738,"score_spread":0.23305255324670923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2932400032","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20415321,0.00030362335,0.7858463,0.00019690448,0.00001931587,0.00028995328,0.00062599604,0.0015040262,0.0070607476],"genre_scores_gemma":[0.77709603,0.00017295926,0.21976638,0.000063137035,0.000037528553,0.00033465264,0.001105555,0.00020261484,0.0012211632],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971373,0.00042097666,0.00038237954,0.0005810024,0.0010979398,0.00038027848],"domain_scores_gemma":[0.98935854,0.004020739,0.0020826948,0.0024602558,0.0018237939,0.00025393962],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016176177,0.0007920959,0.0006400553,0.006843217,0.0008861084,0.0016298728,0.0011011106,0.0010223945,0.0022152448],"category_scores_gemma":[0.013500752,0.00039241606,0.0016061303,0.003841447,0.0020601382,0.003096229,0.0015331097,0.0011308128,0.0003796143],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042579792,0.00060390186,0.07375174,0.00054641144,0.00013736016,0.0012642932,0.0017286235,0.16902556,0.040923283,0.3059751,0.0036799558,0.401938],"study_design_scores_gemma":[0.000042890468,0.0004058591,0.021686845,0.00020370503,0.00015268331,0.0019651374,0.0007321023,0.75123453,0.023211885,0.18555155,0.014680572,0.00013217735],"about_ca_topic_score_codex":0.0047406303,"about_ca_topic_score_gemma":0.0024074225,"teacher_disagreement_score":0.006843217,"about_ca_system_score_codex":0.0013211658,"about_ca_system_score_gemma":0.0013646141,"threshold_uncertainty_score":0.009585798},"labels":[],"label_agreement":null},{"id":"W2934438104","doi":"10.48550/arxiv.1904.01001","title":"Estimation and Prediction of technical debt: a proposal","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University","funders":"","keywords":"Technical debt; Debt; Bankruptcy; Computer science; Payment; Software; Relevance (law); Risk analysis (engineering); Software development; Business; Finance; World Wide Web","score_opus":0.0417757962599083,"score_gpt":0.19530286995385618,"score_spread":0.15352707369394789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2934438104","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0482672,0.013582441,0.9001219,0.016973052,0.0011443235,0.00079209544,0.0019241872,0.0021859426,0.015008892],"genre_scores_gemma":[0.5429418,0.020599147,0.41082853,0.0017271667,0.0029969853,0.0015684619,0.0051338845,0.00024568258,0.013958452],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9960926,0.0012513587,0.0002918212,0.0011766329,0.00089936046,0.00028818627],"domain_scores_gemma":[0.9849626,0.00756743,0.0017550031,0.0011423329,0.003935559,0.00063710107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048186188,0.002043004,0.0013011693,0.005838502,0.000747475,0.0058703595,0.0027274357,0.0027013945,0.0048495824],"category_scores_gemma":[0.02263796,0.0008765854,0.0017669778,0.005636577,0.0011576236,0.007893239,0.0028021766,0.0029237447,0.0027942294],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002990093,0.0011944391,0.14472556,0.0015854909,0.00043487677,0.0010881555,0.0012787312,0.08972784,0.0026379765,0.07088711,0.031407055,0.6547338],"study_design_scores_gemma":[0.00006237518,0.00034510784,0.02881094,0.00090342964,0.00028495668,0.0007230526,0.0013353805,0.84995776,0.0018806057,0.08449386,0.030953713,0.00024880568],"about_ca_topic_score_codex":0.008675899,"about_ca_topic_score_gemma":0.0027022823,"teacher_disagreement_score":0.008675899,"about_ca_system_score_codex":0.0013765928,"about_ca_system_score_gemma":0.0032546967,"threshold_uncertainty_score":0.025483608},"labels":[],"label_agreement":null},{"id":"W2935521157","doi":"10.29019/enfoqueute.v10n1.372","title":"Modelo para estimar el esfuerzo que demanda la automatización de procesos de negocio","year":2019,"lang":"es","type":"article","venue":"Enfoque UTE","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Precision Nanosystems (Canada)","funders":"","keywords":"Humanities; Physics; Philosophy","score_opus":0.010879217336036867,"score_gpt":0.2969737839179239,"score_spread":0.28609456658188703,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2935521157","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.070995934,0.0008633599,0.91321534,0.0008239287,0.00010339442,0.00014989321,0.0008596606,0.0015977985,0.011390657],"genre_scores_gemma":[0.79770625,0.001585355,0.17973638,0.0002001684,0.000065096814,0.0005218101,0.0012438883,0.00035084074,0.018590255],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998781,0.00026657863,0.00007152452,0.00034169247,0.00044128057,0.00009784926],"domain_scores_gemma":[0.9978726,0.001091094,0.00024553854,0.00019255123,0.00053564244,0.000062616054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014313639,0.0012155315,0.00096959586,0.0015057478,0.0005451483,0.0029019434,0.0014615093,0.0015740478,0.005555565],"category_scores_gemma":[0.007252526,0.0007188812,0.00132415,0.0015867981,0.00065649516,0.0028227216,0.0011331345,0.0013892376,0.0014506682],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038501451,0.0001924723,0.017786888,0.0005662046,0.0002459855,0.0001902194,0.00066087785,0.78898937,0.016600484,0.04856677,0.003239967,0.122575775],"study_design_scores_gemma":[0.00002160124,0.00009980276,0.0044414266,0.00007115997,0.000080055455,0.000070528265,0.0001575516,0.96959466,0.0043132124,0.01583192,0.0052757254,0.000042370404],"about_ca_topic_score_codex":0.019029519,"about_ca_topic_score_gemma":0.011879926,"teacher_disagreement_score":0.019029519,"about_ca_system_score_codex":0.0023897474,"about_ca_system_score_gemma":0.0024702933,"threshold_uncertainty_score":0.037837505},"labels":[],"label_agreement":null},{"id":"W2940499664","doi":"10.1109/tse.2019.2912962","title":"The Mutation and Injection Framework: Evaluating Clone Detection Tools with Mutation Analysis","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Computer science; Java; Benchmark (surveying); Mutation; Precision and recall; Mutation testing; Programming language; Data mining; Artificial intelligence; Genetics; Biology; Gene","score_opus":0.014663928079709844,"score_gpt":0.26347687219838184,"score_spread":0.248812944118672,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2940499664","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5235961,0.0025166513,0.42416802,0.000444264,0.00020224809,0.0013419758,0.0019476702,0.041070763,0.004712358],"genre_scores_gemma":[0.5910744,0.0003215819,0.40239242,0.00022199548,0.00004515392,0.00052524154,0.0032743025,0.0011025344,0.0010424269],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98024565,0.0059648217,0.0017276466,0.0029023017,0.008373496,0.0007860214],"domain_scores_gemma":[0.96493495,0.017439043,0.005249152,0.0047995597,0.006593419,0.0009838006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013809586,0.0026109682,0.0013727602,0.010633336,0.0008772677,0.0024567728,0.0036039215,0.0029190304,0.0008057837],"category_scores_gemma":[0.048152328,0.00065107987,0.0018591292,0.003790109,0.0017145191,0.0039145933,0.0025326165,0.0015805727,0.00045868155],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017874695,0.0024105401,0.10266031,0.0023219415,0.001694006,0.0007079409,0.001539988,0.20177193,0.084197946,0.012250714,0.010927849,0.5777294],"study_design_scores_gemma":[0.0002169547,0.0028302567,0.025759598,0.00018491788,0.00028770103,0.0007466789,0.0003063138,0.88250834,0.07660377,0.0039696535,0.0063597662,0.00022611952],"about_ca_topic_score_codex":0.008690738,"about_ca_topic_score_gemma":0.0056877374,"teacher_disagreement_score":0.013809586,"about_ca_system_score_codex":0.00259093,"about_ca_system_score_gemma":0.0029939546,"threshold_uncertainty_score":0.073032975},"labels":[],"label_agreement":null},{"id":"W2944738881","doi":"10.1007/s10664-019-09704-x","title":"cregit: Token-level blame information in git version control repositories","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Victoria","funders":"","keywords":"Blame; Commit; Security token; Computer science; Source code; Code (set theory); Computer security; Programming language; Psychology; Social psychology; Set (abstract data type); Database","score_opus":0.012995523476119583,"score_gpt":0.24247411488059042,"score_spread":0.22947859140447086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2944738881","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05015705,0.0005438632,0.12287758,0.0013330362,0.00037306378,0.00036618562,0.068323456,0.7428213,0.0132044945],"genre_scores_gemma":[0.48187178,0.0005806695,0.1477177,0.00081672583,0.00033913885,0.0005057821,0.27904165,0.0750179,0.014108642],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.99364716,0.0011562118,0.0008388964,0.0010623957,0.0027053033,0.00059000036],"domain_scores_gemma":[0.9535688,0.012430215,0.003786698,0.024117386,0.004576163,0.0015206257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069120456,0.001490458,0.0009179097,0.007960364,0.0010047528,0.0048786066,0.002893015,0.0023790663,0.016584583],"category_scores_gemma":[0.05846504,0.0015566888,0.001055984,0.006626611,0.0010192649,0.011594651,0.0036205226,0.003045109,0.011029558],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0046914606,0.00086189475,0.05192275,0.0016590894,0.00039590008,0.0009510065,0.0024485497,0.016227618,0.012713255,0.02725656,0.47053355,0.41033834],"study_design_scores_gemma":[0.0015723577,0.0011725771,0.05338402,0.0011503906,0.0006094202,0.0023219,0.0008562885,0.26788014,0.1395916,0.08166092,0.44846636,0.0013340526],"about_ca_topic_score_codex":0.0048859804,"about_ca_topic_score_gemma":0.005285226,"teacher_disagreement_score":0.016584583,"about_ca_system_score_codex":0.0015367526,"about_ca_system_score_gemma":0.0025112615,"threshold_uncertainty_score":0.055480957},"labels":[],"label_agreement":null},{"id":"W2945222996","doi":"10.1109/apsec.2018.00054","title":"Why Did This Reviewed Code Crash? An Empirical Study of Mozilla Firefox","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Polytechnique Montréal","funders":"","keywords":"Crash; Computer science; Code review; Code refactoring; Software; Software quality; Software engineering; Code (set theory); Root cause; Software bug; Source lines of code; Software inspection; Empirical research; Computer security; Software development; Reliability engineering; Engineering; Operating system; Programming language","score_opus":0.05654275935009809,"score_gpt":0.3675032591262264,"score_spread":0.3109604997761283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945222996","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9985331,0.00023481475,0.00029733343,0.00017145649,0.000008872394,0.00010072853,0.00012153135,0.000018662266,0.0005133915],"genre_scores_gemma":[0.99718034,0.00030231645,0.0010314434,0.0001986961,0.000021118469,0.00016230346,0.00039698786,0.000038182883,0.0006686255],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98257864,0.007959364,0.0018276166,0.0020729764,0.0047497675,0.00081168744],"domain_scores_gemma":[0.5249991,0.3420792,0.07741385,0.009589222,0.041136596,0.0047820443],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.020572864,0.00050606474,0.0005754906,0.0046849274,0.0020727601,0.0019155991,0.0013008639,0.0013615303,0.0016050899],"category_scores_gemma":[0.20401113,0.00056495116,0.00039453368,0.0030658045,0.002056328,0.0030183913,0.0011804152,0.0017533387,0.0007538376],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009251505,0.0017543226,0.8575037,0.0011520636,0.000191198,0.0022517066,0.08037709,0.0004139682,0.0028056914,0.00055902655,0.0051625655,0.046903577],"study_design_scores_gemma":[0.00010194571,0.0015152466,0.93925947,0.0004939886,0.000114431255,0.0018188432,0.042470466,0.0032959706,0.0017049996,0.000313665,0.008815698,0.000095331015],"about_ca_topic_score_codex":0.008563822,"about_ca_topic_score_gemma":0.013968401,"teacher_disagreement_score":0.97942716,"about_ca_system_score_codex":0.0023305325,"about_ca_system_score_gemma":0.0023901488,"threshold_uncertainty_score":0.10880095},"labels":[],"label_agreement":null},{"id":"W2945475639","doi":"10.1016/j.infsof.2019.05.007","title":"An HMM-based approach for automatic detection and classification of duplicate bug reports","year":2019,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Software regression; Computer science; Software bug; Precision and recall; Hidden Markov model; Software; Rank (graph theory); Crash; Data mining; Information retrieval; Artificial intelligence; Software quality; Software development; Programming language","score_opus":0.010593679754432234,"score_gpt":0.2446091503613171,"score_spread":0.23401547060688488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945475639","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09257045,0.0010570103,0.89033186,0.00032011175,0.00024320185,0.00026614283,0.0018810686,0.0114349965,0.0018951647],"genre_scores_gemma":[0.5676599,0.0006169489,0.42173746,0.00022298469,0.00014994045,0.00027894706,0.0033201613,0.00035622827,0.005657458],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998104,0.0003364125,0.00019118056,0.0005173536,0.0006594343,0.00019158392],"domain_scores_gemma":[0.99507964,0.0018438587,0.00049068633,0.000686411,0.0016851644,0.00021416598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018076788,0.00073920796,0.0013312072,0.003330063,0.00074222166,0.001496863,0.0015125896,0.001516418,0.0017997614],"category_scores_gemma":[0.0052925083,0.00045823885,0.0010143872,0.0023229525,0.00038653,0.0012741224,0.0010921155,0.0013695922,0.001932052],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010066706,0.00057295494,0.033017416,0.0005957547,0.0005126864,0.0007206491,0.00061344146,0.016431833,0.108944595,0.0023314264,0.01064569,0.82460684],"study_design_scores_gemma":[0.000083164654,0.0004502219,0.038116697,0.00010185194,0.00046033564,0.0014907005,0.00032190932,0.8915547,0.053647332,0.0055579795,0.008026667,0.00018843562],"about_ca_topic_score_codex":0.0073890155,"about_ca_topic_score_gemma":0.010005533,"teacher_disagreement_score":0.0073890155,"about_ca_system_score_codex":0.00059197826,"about_ca_system_score_gemma":0.0015813727,"threshold_uncertainty_score":0.0146920085},"labels":[],"label_agreement":null},{"id":"W2945530585","doi":"10.1007/s10664-019-09719-4","title":"Fostering good coding practices through individual feedback and gamification: an industrial case study","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Coding (social sciences); Variety (cybernetics); Computer science; Best practice; Code review; Software; Quality (philosophy); Software engineering; Knowledge management; Software quality; Data science; Software development; Management","score_opus":0.1720305715459125,"score_gpt":0.37070826398116774,"score_spread":0.19867769243525524,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945530585","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9905956,0.000037964255,0.00637731,0.00021076004,0.0000072896655,0.00018220498,0.000013608462,0.000059437163,0.002515754],"genre_scores_gemma":[0.9833059,0.000051763964,0.015321982,0.00003996227,0.0000036339643,0.000099608325,0.000021282884,0.000016286778,0.0011396612],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9950151,0.0031465637,0.0001589203,0.00036743036,0.00077013735,0.00054178486],"domain_scores_gemma":[0.9541145,0.033474654,0.0023590932,0.0046997676,0.002958063,0.0023939894],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010050518,0.0007212798,0.00037374144,0.0013913481,0.0021428291,0.0019307119,0.0019714208,0.0018792866,0.0017233766],"category_scores_gemma":[0.03139336,0.00040394146,0.00032623482,0.0009321265,0.002145197,0.001470084,0.0026925707,0.0015960177,0.00036390754],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003174185,0.0573215,0.13368413,0.00090333406,0.00016036138,0.0063621863,0.10651066,0.032266267,0.022857388,0.016454607,0.004168423,0.616137],"study_design_scores_gemma":[0.0032610483,0.052707862,0.17532249,0.001693288,0.0006325358,0.010524913,0.18989782,0.3676045,0.10715235,0.04738656,0.042923257,0.00089339],"about_ca_topic_score_codex":0.002017097,"about_ca_topic_score_gemma":0.0044750436,"teacher_disagreement_score":0.98994946,"about_ca_system_score_codex":0.0014228699,"about_ca_system_score_gemma":0.0023665843,"threshold_uncertainty_score":0.05315286},"labels":[],"label_agreement":null},{"id":"W2945596753","doi":"10.1109/icsa-c.2019.00023","title":"Component Comparison, Evaluation, and Selection: A Continuous Approach","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Carnegie Mellon University; U.S. Department of Defense","keywords":"Component (thermodynamics); Computer science; Balanced scorecard; Selection (genetic algorithm); Context (archaeology); Software; Agile software development; Software quality; Software engineering; Data mining; Software development; Artificial intelligence; Process management; Engineering","score_opus":0.024955493343050493,"score_gpt":0.28915450971480605,"score_spread":0.2641990163717556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945596753","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03141705,0.0022815012,0.9565413,0.001325796,0.0001247196,0.0012825553,0.00018005582,0.0006981077,0.0061488897],"genre_scores_gemma":[0.2689977,0.00069594535,0.7270945,0.00019935015,0.00014617741,0.0011683642,0.00033845627,0.00013586268,0.0012237238],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.82426775,0.10958335,0.010869266,0.00906707,0.044343375,0.0018693076],"domain_scores_gemma":[0.6664122,0.24543056,0.016926352,0.01958622,0.047894794,0.0037498947],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11569626,0.0025808197,0.0049296417,0.019727968,0.002367306,0.014560142,0.005138202,0.0026169876,0.0040207785],"category_scores_gemma":[0.27208963,0.0013638105,0.00186041,0.015520303,0.006760783,0.011944192,0.0065755188,0.0034004706,0.00074837485],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008435101,0.00080836075,0.028175335,0.0020977042,0.000809928,0.00018940851,0.002868465,0.029002018,0.0030062015,0.08180266,0.004062289,0.8463341],"study_design_scores_gemma":[0.00089303387,0.004593281,0.043282475,0.002041461,0.0009880601,0.0007125239,0.0047953,0.51963913,0.011425509,0.38271314,0.028291214,0.0006248616],"about_ca_topic_score_codex":0.0027094833,"about_ca_topic_score_gemma":0.002268004,"teacher_disagreement_score":0.11569626,"about_ca_system_score_codex":0.0050312225,"about_ca_system_score_gemma":0.006544828,"threshold_uncertainty_score":0.6118676},"labels":[],"label_agreement":null},{"id":"W2945826489","doi":"10.1007/s10664-019-09788-5","title":"MSRBot: Using bots to answer questions from software repositories","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Software; Process (computing); Software development; Work (physics); Code (set theory); Software analytics; Team software process; Source code; Software peer review; Verification and validation","score_opus":0.04104006454653197,"score_gpt":0.29623689857983276,"score_spread":0.2551968340333008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945826489","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42823046,0.0015868532,0.39150453,0.0052291206,0.0010672666,0.002322056,0.017983435,0.12403758,0.028038627],"genre_scores_gemma":[0.6853954,0.000481928,0.27351126,0.0021110484,0.00021204336,0.0014351932,0.017362762,0.0025535745,0.016936846],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946234,0.0025964368,0.0003350834,0.0007430606,0.0014365925,0.00026548133],"domain_scores_gemma":[0.9725859,0.020389713,0.0015356172,0.0028049336,0.001826438,0.0008573316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040606447,0.001219173,0.00079019764,0.003911489,0.0011750847,0.0014091885,0.0014836356,0.0024289144,0.0056175217],"category_scores_gemma":[0.03231308,0.0005253039,0.0005248169,0.0014579347,0.0007297025,0.004845444,0.003626405,0.0014198284,0.00413536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030012142,0.002534155,0.08138831,0.0039001373,0.00058106886,0.0019320701,0.014815794,0.010289581,0.04333743,0.020229934,0.18298322,0.635007],"study_design_scores_gemma":[0.00079879776,0.0023220426,0.0508546,0.00082436454,0.00038585675,0.001536536,0.012039081,0.5880781,0.0373174,0.10580501,0.19965313,0.0003851582],"about_ca_topic_score_codex":0.0048720436,"about_ca_topic_score_gemma":0.010518932,"teacher_disagreement_score":0.0056175217,"about_ca_system_score_codex":0.00087518786,"about_ca_system_score_gemma":0.0013829935,"threshold_uncertainty_score":0.021474957},"labels":[],"label_agreement":null},{"id":"W2945944668","doi":"10.1007/s10515-019-00256-4","title":"Improving web service interfaces modularity using multi-objective optimization","year":2019,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Interface (matter); Code refactoring; Service (business); Web service; Software engineering; Service-oriented architecture; Artifact (error); Reuse; Modular programming; Service provider; World Wide Web; Programming language; Artificial intelligence; Software; Operating system; Engineering","score_opus":0.013282066220827,"score_gpt":0.2464925954551059,"score_spread":0.2332105292342789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945944668","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18816912,0.00024026865,0.80760497,0.00013241249,0.00003937333,0.0000934541,0.000029816278,0.0010059747,0.0026846044],"genre_scores_gemma":[0.8079931,0.00006927171,0.19019724,0.00004791909,0.000014787261,0.0000652391,0.000057078218,0.0002600437,0.0012952297],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992106,0.0002541548,0.000034361754,0.00012639204,0.00023886748,0.00013572616],"domain_scores_gemma":[0.9985972,0.0006016524,0.00026439768,0.00014965232,0.0002937162,0.000093459465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014275447,0.0014052886,0.0010827215,0.001194507,0.0005435484,0.00093806785,0.0011878166,0.00085246755,0.0015989351],"category_scores_gemma":[0.0033942547,0.00051904283,0.00086303154,0.00068394706,0.00060827506,0.0011915755,0.001227163,0.0009536998,0.00029228596],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012009433,0.00022436788,0.0010810525,0.0000705424,0.00006278762,0.00005087888,0.000042411524,0.91941416,0.016702142,0.002488252,0.0003531068,0.05939016],"study_design_scores_gemma":[0.000007482629,0.000030696454,0.00009728723,0.0000023507891,0.000010698467,0.00000702205,0.0000064600704,0.9978078,0.0014169816,0.0005372796,0.00007373886,0.0000022036072],"about_ca_topic_score_codex":0.002629002,"about_ca_topic_score_gemma":0.0026953814,"teacher_disagreement_score":0.002629002,"about_ca_system_score_codex":0.0010437026,"about_ca_system_score_gemma":0.0011730901,"threshold_uncertainty_score":0.007572651},"labels":[],"label_agreement":null},{"id":"W2946233956","doi":"10.1109/tse.2019.2918520","title":"Characterizing Crowds to Better Optimize Worker Recommendation in Crowdsourced Testing","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Key Research and Development Program of China; China Scholarship Council; National Natural Science Foundation of China","keywords":"Computer science; Crowds; Crowdsourcing; Task (project management); Context (archaeology); Software bug; Relevance (law); Test (biology); Machine learning; Software; Data science; Computer security; World Wide Web; Engineering","score_opus":0.01883341979320734,"score_gpt":0.24014374368048746,"score_spread":0.2213103238872801,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2946233956","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16749662,0.0032160548,0.81617427,0.0013169921,0.00027263697,0.0008239317,0.00088982296,0.004227981,0.005581686],"genre_scores_gemma":[0.8155667,0.0006038613,0.17698951,0.0006349104,0.00017905585,0.0005656869,0.0013756427,0.00031602627,0.0037685204],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9970777,0.0009854118,0.00015636731,0.0009774049,0.00054352987,0.00025950721],"domain_scores_gemma":[0.99048525,0.005975623,0.00079924625,0.0009795956,0.001122556,0.0006377077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034388425,0.002437201,0.0026087153,0.0019063995,0.0011175921,0.0015434986,0.0032568104,0.0021177912,0.0021235915],"category_scores_gemma":[0.016308047,0.00091825705,0.0011010173,0.0016038334,0.0009977957,0.0022367279,0.0019158508,0.0012528821,0.0010752406],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012775373,0.0011055141,0.026799142,0.0011001528,0.00040802415,0.0005775124,0.0010129936,0.49781248,0.013002403,0.003870812,0.014634171,0.43839926],"study_design_scores_gemma":[0.00012101016,0.00024384398,0.0028614448,0.00005864278,0.000086518696,0.00010407621,0.0002612292,0.984938,0.0020773546,0.0061388835,0.0030654576,0.00004349176],"about_ca_topic_score_codex":0.016239008,"about_ca_topic_score_gemma":0.018285027,"teacher_disagreement_score":0.016239008,"about_ca_system_score_codex":0.0012552363,"about_ca_system_score_gemma":0.0022207724,"threshold_uncertainty_score":0.03228897},"labels":[],"label_agreement":null},{"id":"W2946377044","doi":"10.1109/icse-companion.2019.00121","title":"Analyzing and Repairing Compilation Errors","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Symbol (formal); Java; Set (abstract data type); Syntax; Tree (set theory); Abstract syntax tree; Resolution (logic); Artificial intelligence; Machine translation; Software engineering; Programming language; Machine learning","score_opus":0.014950139281644317,"score_gpt":0.2598785073467029,"score_spread":0.2449283680650586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2946377044","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9067645,0.00083013,0.07499912,0.0005304342,0.00009454447,0.000066622815,0.0010406971,0.013495381,0.0021785633],"genre_scores_gemma":[0.9366697,0.00028619476,0.05888306,0.00008779916,0.000026120493,0.000033140346,0.0017578783,0.0007075125,0.0015484769],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99716634,0.00078519667,0.0002261954,0.0006920272,0.00088422204,0.00024604757],"domain_scores_gemma":[0.98149806,0.010006357,0.0029100839,0.0025201784,0.0028579102,0.00020736165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024055475,0.00086453033,0.00056766893,0.0027678104,0.00047854145,0.0010451075,0.000784571,0.0008048452,0.0008648538],"category_scores_gemma":[0.02680141,0.0005017058,0.0005344893,0.0014047694,0.0005497817,0.0015294848,0.00074760104,0.0009171768,0.0007694311],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038427406,0.0002863094,0.22703241,0.0007536294,0.00018940125,0.0011724767,0.0020599272,0.07505729,0.029650396,0.0016309865,0.0087358635,0.653047],"study_design_scores_gemma":[0.00004043735,0.00040004545,0.14881234,0.00022315477,0.0002928554,0.0012687776,0.0012970129,0.7519659,0.07283827,0.007628652,0.015108808,0.00012379966],"about_ca_topic_score_codex":0.006389138,"about_ca_topic_score_gemma":0.007980067,"teacher_disagreement_score":0.006389138,"about_ca_system_score_codex":0.00058118254,"about_ca_system_score_gemma":0.0012381832,"threshold_uncertainty_score":0.012721956},"labels":[],"label_agreement":null},{"id":"W2948724728","doi":"10.1109/icpc.2019.00053","title":"On the Use of Information Retrieval to Automate the Detection of Third-Party Java Library Migration at the Method Level","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Java; Benchmark (surveying); Process (computing); Documentation; Context (archaeology); Matching (statistics); Information retrieval; Plug-in; Similarity (geometry); Data mining; Scale (ratio); Artificial intelligence; Programming language; Image (mathematics)","score_opus":0.06556379480234045,"score_gpt":0.2919367685898256,"score_spread":0.22637297378748514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2948724728","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1366006,0.0048509794,0.80506855,0.0015220732,0.00023240814,0.0014277993,0.0071773683,0.03298439,0.010135858],"genre_scores_gemma":[0.19553114,0.0013660705,0.7886917,0.0005365594,0.00008138073,0.0004836328,0.009634844,0.00061664265,0.0030580235],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950924,0.00091940496,0.0006011863,0.0012984831,0.0018436923,0.00024486516],"domain_scores_gemma":[0.98375875,0.006662779,0.0028698617,0.0028257743,0.0035749932,0.0003078416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004145143,0.0013743404,0.0012691022,0.021707376,0.0012520619,0.0025035548,0.002299259,0.0016846355,0.0014250218],"category_scores_gemma":[0.018059246,0.00056187785,0.0015100057,0.014143648,0.00073290063,0.0035650337,0.0021484043,0.0012083057,0.0027368197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026009406,0.00071069057,0.03954657,0.0015573961,0.00029681277,0.0006536332,0.001220281,0.005075486,0.037809864,0.004232527,0.01869301,0.88994354],"study_design_scores_gemma":[0.00017254373,0.0008496384,0.116999954,0.0012147097,0.0007405344,0.0047941986,0.002605871,0.57552314,0.15562443,0.032733098,0.10825546,0.0004864182],"about_ca_topic_score_codex":0.009232729,"about_ca_topic_score_gemma":0.01543136,"teacher_disagreement_score":0.021707376,"about_ca_system_score_codex":0.0010088764,"about_ca_system_score_gemma":0.003116928,"threshold_uncertainty_score":0.021921873},"labels":[],"label_agreement":null},{"id":"W2949142318","doi":"10.48550/arxiv.1807.02274","title":"Recommending Relevant Sections from a Webpage about Programming Errors\\n and Exceptions","year":2018,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Web page; Computer science; Context (archaeology); World Wide Web; Information retrieval; Software; The Internet; Precision and recall; Page view; Static web page; Web development; Programming language","score_opus":0.08752907531005548,"score_gpt":0.22317878461442045,"score_spread":0.13564970930436498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949142318","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5504065,0.012931839,0.32681566,0.0022398115,0.0010060354,0.0019719563,0.011136635,0.06371412,0.029777365],"genre_scores_gemma":[0.48435548,0.004590812,0.4636665,0.00047151482,0.00044308917,0.00039066025,0.018053513,0.0012252263,0.026803216],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999559,0.00006320874,0.000037935206,0.00012418612,0.0001858643,0.000029815497],"domain_scores_gemma":[0.99769765,0.00089487666,0.0002758927,0.00025361634,0.0007004719,0.0001774515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00041006028,0.0012131337,0.0006779391,0.0045290403,0.0006005172,0.0010761551,0.00074167794,0.0009351515,0.004043635],"category_scores_gemma":[0.004490876,0.0005099285,0.0007780682,0.0022208907,0.00021163942,0.001567723,0.00045120652,0.00082123245,0.0041548745],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005194112,0.00072722964,0.034711134,0.00164043,0.00020280863,0.0006427941,0.00070006255,0.004555351,0.041081905,0.0009763163,0.052946683,0.86129594],"study_design_scores_gemma":[0.00025796492,0.0019336563,0.15585393,0.0012715092,0.001809893,0.005245436,0.0030765317,0.4439585,0.1451699,0.008760158,0.2322854,0.00037715648],"about_ca_topic_score_codex":0.0063772798,"about_ca_topic_score_gemma":0.018694256,"teacher_disagreement_score":0.0063772798,"about_ca_system_score_codex":0.0002877613,"about_ca_system_score_gemma":0.0013059169,"threshold_uncertainty_score":0.013527334},"labels":[],"label_agreement":null},{"id":"W2950236608","doi":"10.1007/s10664-019-09709-6","title":"A study of build inflation in 30 million CPAN builds on 13 Perl versions and 10 operating systems","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Polytechnique Montréal","funders":"","keywords":"Perl; Inflation (cosmology); Computer science; Operating system; Programming language; Physics; Astronomy","score_opus":0.021393192071352304,"score_gpt":0.28106298312999833,"score_spread":0.259669791058646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950236608","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9988696,0.000037575635,0.00009544048,0.00007244752,0.0000022320237,0.000003141552,0.0002163593,0.000011188925,0.0006919959],"genre_scores_gemma":[0.99917245,0.000017626264,0.000058733935,0.0000144414935,0.0000051852267,0.0000029322105,0.00044471247,0.0000054639927,0.00027830937],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9985839,0.0004068725,0.00008286235,0.0002061501,0.00046967826,0.00025049003],"domain_scores_gemma":[0.97516114,0.013136293,0.006395646,0.0008895278,0.0031963123,0.0012210915],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.001766184,0.00023069514,0.00026981026,0.0013098102,0.00063455256,0.0011018997,0.0006418735,0.00063866226,0.0016205243],"category_scores_gemma":[0.015758624,0.00037020585,0.00036734328,0.002509547,0.00056436076,0.0011835583,0.0006348196,0.0015616645,0.00041736092],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088094064,0.00063267845,0.9737879,0.000029638975,0.00013457352,0.00037784342,0.0012564866,0.0059433845,0.000739736,0.0012961759,0.0020594287,0.012861221],"study_design_scores_gemma":[0.0000138450605,0.00023605287,0.9840206,0.000007914457,0.000041366093,0.00014764628,0.0014397546,0.012426332,0.00047182105,0.00034879535,0.0008244439,0.000021483747],"about_ca_topic_score_codex":0.029236505,"about_ca_topic_score_gemma":0.02771707,"teacher_disagreement_score":0.9982338,"about_ca_system_score_codex":0.0019864342,"about_ca_system_score_gemma":0.0007186261,"threshold_uncertainty_score":0.05813265},"labels":[],"label_agreement":null},{"id":"W2950559347","doi":"10.48550/arxiv.1703.04564","title":"Analogy-based effort estimation: a new method to discover set of analogies from dataset characteristics","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Analogy; Computer science; Set (abstract data type); Cluster analysis; Data mining; Medoid; Estimation; Machine learning; Artificial intelligence; Engineering","score_opus":0.11629210779738773,"score_gpt":0.2834634239446617,"score_spread":0.16717131614727393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950559347","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020988453,0.00024686242,0.97571814,0.0001701103,0.00003104451,0.0001661938,0.0004246547,0.0010958898,0.0011586434],"genre_scores_gemma":[0.38272864,0.0002966618,0.6125544,0.00016346636,0.00008544544,0.0006405116,0.0017204651,0.00017932153,0.0016311227],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99638367,0.0010029124,0.0003321972,0.001187719,0.0009538986,0.00013955493],"domain_scores_gemma":[0.99043965,0.005072847,0.0015184232,0.0013739623,0.0013548029,0.0002403159],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002843587,0.0013112556,0.0012594956,0.006457815,0.00067693193,0.0016477917,0.0019317458,0.0013895399,0.0022484965],"category_scores_gemma":[0.023397267,0.0005603055,0.0016431123,0.004671036,0.00081935845,0.0041620918,0.0022141826,0.001875474,0.0007538728],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005179465,0.00045086295,0.025904102,0.00057814317,0.0005230929,0.00024814505,0.00096957694,0.19895147,0.00801666,0.02021343,0.005164835,0.7384618],"study_design_scores_gemma":[0.000034589983,0.00019951249,0.006542231,0.000055675177,0.000058798654,0.00020689075,0.00016455897,0.95786333,0.0025592053,0.028278248,0.003969524,0.00006746084],"about_ca_topic_score_codex":0.0025654447,"about_ca_topic_score_gemma":0.002810204,"teacher_disagreement_score":0.006457815,"about_ca_system_score_codex":0.001080951,"about_ca_system_score_gemma":0.0012180244,"threshold_uncertainty_score":0.01503849},"labels":[],"label_agreement":null},{"id":"W2950609722","doi":"10.48550/arxiv.1808.00594","title":"Improving IR-Based Bug Localization with Context-Aware Query Reformulation","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Context (archaeology); Query expansion; Baseline (sea); Query language; State (computer science); Data mining; Programming language","score_opus":0.04577286314861694,"score_gpt":0.18931150841012168,"score_spread":0.14353864526150473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950609722","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10692705,0.009344065,0.78475547,0.0016556607,0.00041050182,0.0007082162,0.0020286487,0.08980401,0.004366503],"genre_scores_gemma":[0.3148223,0.0018624843,0.67015654,0.0009603562,0.00041971385,0.00026659027,0.0054116542,0.0018528077,0.004247563],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952077,0.0012547801,0.00045407933,0.0011151235,0.0016560508,0.00031220962],"domain_scores_gemma":[0.99052936,0.003954993,0.0010485242,0.0018874641,0.0023813169,0.00019825599],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030913064,0.0027718276,0.0027072616,0.006801641,0.0008831982,0.0018104721,0.0023349246,0.0016778741,0.0037505946],"category_scores_gemma":[0.017046787,0.0006252689,0.0019966129,0.0038577355,0.00091180886,0.0044608647,0.0023105103,0.0018284164,0.003657816],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058167614,0.0005195577,0.0056208638,0.0016111402,0.00020152422,0.0005995407,0.0011836998,0.013435718,0.10706409,0.0036309543,0.03513845,0.83041286],"study_design_scores_gemma":[0.0005136665,0.0018547258,0.011972514,0.00023824508,0.0013637078,0.003743283,0.001570053,0.696799,0.210736,0.014845977,0.055967625,0.00039517137],"about_ca_topic_score_codex":0.008364242,"about_ca_topic_score_gemma":0.006062548,"teacher_disagreement_score":0.008364242,"about_ca_system_score_codex":0.001063136,"about_ca_system_score_gemma":0.00211203,"threshold_uncertainty_score":0.016631067},"labels":[],"label_agreement":null},{"id":"W2951052716","doi":"10.7287/peerj.preprints.3186","title":"Improved query reformulation for concept location using CodeRank and document structures","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Query expansion; Information retrieval; Source code; Query optimization; Sargable; Web query classification; Baseline (sea); Web search query; Software; Task (project management); Code (set theory); Term (time); Quality (philosophy); Query language; Data mining; Search engine; Programming language","score_opus":0.036999784354417545,"score_gpt":0.3369293061214571,"score_spread":0.2999295217670396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951052716","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07810368,0.0022974003,0.89267164,0.0011476004,0.00023003656,0.0009454075,0.0023546354,0.01981631,0.00243324],"genre_scores_gemma":[0.2402604,0.00071382715,0.74500746,0.0003232475,0.0002788995,0.00044324948,0.008361462,0.00081918103,0.003792221],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9944482,0.0019343595,0.0005653404,0.0008172839,0.0019356642,0.00029919416],"domain_scores_gemma":[0.98721826,0.0066015166,0.00093008886,0.0019402175,0.0030587972,0.00025117776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031766044,0.0017271751,0.002143031,0.006207652,0.0010727216,0.0021614907,0.0018744958,0.0013543746,0.0042230627],"category_scores_gemma":[0.019568896,0.0005140199,0.0013582044,0.004811225,0.0009397527,0.005504767,0.0020104004,0.0021002435,0.0027814223],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008566056,0.0007436592,0.003302186,0.0011394573,0.0001548265,0.00041668268,0.0014171717,0.03821396,0.077508144,0.014999439,0.03952463,0.8217233],"study_design_scores_gemma":[0.00036840927,0.0007873842,0.002751858,0.00007320365,0.00020286253,0.0010978668,0.0011682402,0.86269146,0.0806636,0.016086891,0.033919718,0.000188524],"about_ca_topic_score_codex":0.010547712,"about_ca_topic_score_gemma":0.010044632,"teacher_disagreement_score":0.010547712,"about_ca_system_score_codex":0.0015044564,"about_ca_system_score_gemma":0.0031490987,"threshold_uncertainty_score":0.02097261},"labels":[],"label_agreement":null},{"id":"W2951710749","doi":"10.1109/tse.2019.2924006","title":"Locating Latent Design Information in Developer Discussions: A Study on Pull Requests","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Université de Montréal; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Maintainability; Classifier (UML); Documentation; Software engineering; Software design; Machine learning; Software; Robustness (evolution); Source lines of code; Artificial intelligence; Data mining; Software development; Programming language","score_opus":0.021608475185581517,"score_gpt":0.24979844873978285,"score_spread":0.22818997355420134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951710749","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.994009,0.00018931473,0.003926924,0.00018283108,0.000010074053,0.000060034192,0.00023087517,0.000107106345,0.0012838924],"genre_scores_gemma":[0.99321884,0.0001466108,0.003761887,0.00013154709,0.000032260745,0.00010732825,0.0010035459,0.000099012985,0.001499072],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9900956,0.0059546134,0.0007297847,0.00095446245,0.0018936532,0.00037184317],"domain_scores_gemma":[0.7545282,0.20457844,0.017771086,0.0067908634,0.014059922,0.0022714643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009141678,0.00045879395,0.00052441284,0.004007635,0.0016352152,0.0017757871,0.0008543476,0.0014571769,0.0013321654],"category_scores_gemma":[0.09755081,0.00036611754,0.00041472173,0.0025459128,0.0010535562,0.0040167435,0.0018201839,0.0016531975,0.00086731795],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014355687,0.0022477326,0.74849766,0.0008249818,0.000135155,0.0015044395,0.094988994,0.0012562734,0.016018867,0.001511595,0.004763383,0.12681545],"study_design_scores_gemma":[0.00013750218,0.0021606428,0.8552497,0.00036682285,0.00014957282,0.0027406942,0.057276353,0.04226146,0.012826707,0.003572237,0.023058092,0.0002002378],"about_ca_topic_score_codex":0.002327137,"about_ca_topic_score_gemma":0.0038675515,"teacher_disagreement_score":0.009141678,"about_ca_system_score_codex":0.0009655558,"about_ca_system_score_gemma":0.0005779243,"threshold_uncertainty_score":0.0483464},"labels":[],"label_agreement":null},{"id":"W2952129634","doi":"10.1109/se4science.2019.00011","title":"Debunking the Myth That Upfront Requirements Are Infeasible for Scientific Computing Software","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software engineering; Documentation; Software requirements; Software requirements specification; Requirements analysis; Traceability; Requirements traceability; Software; Software development; Software design; Systems engineering; Requirement; Programming language; Engineering","score_opus":0.1016743850693255,"score_gpt":0.3245010695131971,"score_spread":0.22282668444387163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952129634","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022042306,0.008381595,0.2933673,0.6081439,0.0057899747,0.00006954506,0.00013760602,0.001436769,0.06063104],"genre_scores_gemma":[0.49992335,0.01658498,0.30407742,0.13116899,0.009095007,0.00050803396,0.0002657057,0.0030761715,0.035300408],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97655404,0.012037821,0.0010924255,0.0022671716,0.0073458008,0.00070281676],"domain_scores_gemma":[0.8880368,0.07871795,0.0036064868,0.017139507,0.0107690655,0.0017301562],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028774777,0.00086950033,0.00070153165,0.0016566362,0.0042436155,0.009213734,0.002646488,0.0065175234,0.006024415],"category_scores_gemma":[0.08706245,0.0009045187,0.00096577295,0.0009460921,0.034652222,0.031216271,0.0065389895,0.023529934,0.002898838],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000074408454,0.000044383083,0.0005251355,0.0002224233,0.000030128402,0.0002069063,0.008710276,0.0017199772,0.0011111364,0.8862294,0.051789075,0.049336858],"study_design_scores_gemma":[0.000029009938,0.0000490241,0.00023808563,0.0003338196,0.000012633044,0.0003347765,0.0016778582,0.0024788743,0.0010106618,0.79773515,0.19604364,0.000056428817],"about_ca_topic_score_codex":0.0018298503,"about_ca_topic_score_gemma":0.001288912,"teacher_disagreement_score":0.9712252,"about_ca_system_score_codex":0.0038218612,"about_ca_system_score_gemma":0.0033622,"threshold_uncertainty_score":0.1521774},"labels":[],"label_agreement":null},{"id":"W2952525141","doi":"10.48550/arxiv.1807.04479","title":"RACK: Code Search in the IDE using Crowdsourced Knowledge","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Information retrieval; Web search query; Code (set theory); Source code; Context (archaeology); Matching (statistics); Programming language; Query expansion; Search engine; World Wide Web; Database; Set (abstract data type)","score_opus":0.15997904129413945,"score_gpt":0.25466046705107626,"score_spread":0.0946814257569368,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952525141","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07716726,0.002128215,0.6170535,0.002767477,0.000537043,0.0026249918,0.0571985,0.19533496,0.045188077],"genre_scores_gemma":[0.2704335,0.00075790193,0.6387317,0.00082476315,0.00013534991,0.0015829005,0.07010081,0.0061799274,0.011253171],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956583,0.0012769692,0.00028985986,0.0012376668,0.0012718512,0.00026538325],"domain_scores_gemma":[0.99029386,0.005341075,0.00051251583,0.0023664662,0.0009765937,0.00050950696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040277867,0.0017544607,0.0012897631,0.007299864,0.0013204883,0.0026365346,0.0022985602,0.0017239108,0.0071132304],"category_scores_gemma":[0.018024197,0.00053754146,0.0013108992,0.0036742836,0.0010758564,0.0047766296,0.006475897,0.0013015541,0.0070481673],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018729101,0.001006834,0.012873226,0.0031809546,0.00046487735,0.001687876,0.004158527,0.02796208,0.021168564,0.02769603,0.26132414,0.636604],"study_design_scores_gemma":[0.0006966861,0.0004141231,0.008876092,0.0004713524,0.00016317738,0.0007282782,0.0046153455,0.58431536,0.032560006,0.116730615,0.2500405,0.00038838745],"about_ca_topic_score_codex":0.009800304,"about_ca_topic_score_gemma":0.014579063,"teacher_disagreement_score":0.009800304,"about_ca_system_score_codex":0.001161489,"about_ca_system_score_gemma":0.0025617054,"threshold_uncertainty_score":0.023796082},"labels":[],"label_agreement":null},{"id":"W2952939633","doi":"","title":"A Neuro-Fuzzy Model for Function Point Calibration","year":2015,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Bank of Canada","funders":"","keywords":"Computer science; Fuzzy logic; Calibration; Software; Function point; Artificial neural network; Data mining; Neuro-fuzzy; Artificial intelligence; Machine learning; Point (geometry); Function (biology); Fuzzy control system; Software development; Mathematics; Statistics","score_opus":0.06231095103633173,"score_gpt":0.29029115694335844,"score_spread":0.22798020590702672,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952939633","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0109671205,0.00034697878,0.98070383,0.00027018867,0.00008670077,0.000052269283,0.00015800208,0.0003486652,0.007066368],"genre_scores_gemma":[0.81453586,0.00081072043,0.16439891,0.00015144375,0.00008988512,0.00035889508,0.00036379756,0.00010918543,0.019181322],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99954396,0.00010281695,0.000023223327,0.00013306302,0.00015451368,0.00004238648],"domain_scores_gemma":[0.99934417,0.00025043846,0.000067497094,0.000064736734,0.00025294677,0.000020228148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009604069,0.00077894103,0.0007994004,0.00078952545,0.00066126656,0.0016842225,0.001909974,0.0020429445,0.0044375807],"category_scores_gemma":[0.0030594044,0.00046237366,0.000899036,0.0010900783,0.0006892177,0.0015178359,0.0006572276,0.0019359723,0.0015476015],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034579298,0.000020915517,0.0003511972,0.000040961062,0.000023404884,0.000057896552,0.000053967324,0.9649589,0.0011239324,0.0110809775,0.0007331623,0.021520026],"study_design_scores_gemma":[0.000002604117,0.000009647477,0.00013383017,0.000007990156,0.00000504582,0.00001730505,0.000004692424,0.99579054,0.00026129186,0.003057264,0.0007022934,0.0000073882206],"about_ca_topic_score_codex":0.021096295,"about_ca_topic_score_gemma":0.012827063,"teacher_disagreement_score":0.021096295,"about_ca_system_score_codex":0.0014509527,"about_ca_system_score_gemma":0.0011798177,"threshold_uncertainty_score":0.041947007},"labels":[],"label_agreement":null},{"id":"W2953266794","doi":"","title":"Enhancing Use Case Points Estimation Method Using Soft Computing Techniques","year":2016,"lang":"en","type":"preprint","venue":"Scholarship@Western (Western University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Use Case Points; Software sizing; Software metric; Computer science; Estimation; Software; Software development; Metric (unit); Schedule; Software construction; Reliability engineering; Data mining; Systems engineering; Engineering","score_opus":0.13661127326781436,"score_gpt":0.3816669547789311,"score_spread":0.24505568151111673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953266794","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011262772,0.000087766966,0.9869348,0.000054813125,0.000020281534,0.00006927356,0.000026623698,0.00053063233,0.0010130748],"genre_scores_gemma":[0.34008947,0.0002927261,0.65679574,0.00005444144,0.00004149253,0.00022957895,0.0002092933,0.00012155181,0.0021656197],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968035,0.0007193155,0.00027274498,0.000403709,0.0016507086,0.00015003858],"domain_scores_gemma":[0.9914323,0.0040728436,0.0008144211,0.0006315219,0.002909664,0.00013917073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023101226,0.0011967364,0.0010944518,0.005206505,0.0006405758,0.0018379864,0.0015062769,0.0012579988,0.0035606287],"category_scores_gemma":[0.011903805,0.0006503177,0.0012460818,0.0028548858,0.0004634075,0.002283948,0.0012096126,0.0016084388,0.0010477368],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021103953,0.00027970434,0.0045795348,0.00034368626,0.00012479459,0.00024320888,0.00041695387,0.21288233,0.027302938,0.0076342076,0.0015397054,0.74444187],"study_design_scores_gemma":[0.000010137665,0.00004493063,0.0010997896,0.000027281892,0.000029407305,0.00008879855,0.00006710728,0.98374337,0.010522528,0.0026929912,0.0016488791,0.000024730438],"about_ca_topic_score_codex":0.004498021,"about_ca_topic_score_gemma":0.0035570913,"teacher_disagreement_score":0.005206505,"about_ca_system_score_codex":0.00065526343,"about_ca_system_score_gemma":0.0010003971,"threshold_uncertainty_score":0.012217283},"labels":[],"label_agreement":null},{"id":"W2953336937","doi":"10.1002/smr.2180","title":"Analysis of cluster center initialization of 2FA‐kprototypes analogy‐based software effort estimation","year":2019,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Initialization; Computer science; Unpacking; Cluster analysis; Cluster (spacecraft); Analogy; Fuzzy logic; Categorical variable; Software; Artificial intelligence; Machine learning; Operating system","score_opus":0.00986538643137731,"score_gpt":0.27792727701311865,"score_spread":0.26806189058174135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953336937","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6452189,0.00049326837,0.34724152,0.0001748789,0.00011452805,0.00017621116,0.00096908194,0.003415166,0.0021964496],"genre_scores_gemma":[0.9131338,0.000054327837,0.084327705,0.000023143184,0.000012024036,0.00008823645,0.0016376855,0.00008711105,0.0006359474],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99797577,0.00058562536,0.00013003673,0.0005363117,0.00055885845,0.0002133549],"domain_scores_gemma":[0.9919503,0.0036603322,0.0007016684,0.00089578354,0.0025788725,0.00021296497],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028342837,0.0010171153,0.0007055108,0.0031748323,0.00063104223,0.0010097544,0.00130067,0.0008378653,0.0008755186],"category_scores_gemma":[0.016242102,0.00022629299,0.0006491161,0.0018466619,0.00036981114,0.0012080357,0.00091458834,0.0006568084,0.00043160727],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016137558,0.00035687294,0.10411669,0.00029967545,0.00026868522,0.00028031718,0.0006069218,0.4839438,0.0147630535,0.005401996,0.006337855,0.38201046],"study_design_scores_gemma":[0.00001631083,0.000086172535,0.017324451,0.00001977077,0.000023420629,0.000050231225,0.00012621276,0.97186434,0.008728857,0.00078323775,0.0009471223,0.000029891802],"about_ca_topic_score_codex":0.01607969,"about_ca_topic_score_gemma":0.010607412,"teacher_disagreement_score":0.01607969,"about_ca_system_score_codex":0.0012172868,"about_ca_system_score_gemma":0.00094455696,"threshold_uncertainty_score":0.03197217},"labels":[],"label_agreement":null},{"id":"W2953431343","doi":"10.1109/icse-companion.2019.00039","title":"Witt: Querying Technology Terms Based on Automated Classification","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Categorization; Software; Term (time); Java; Spurious relationship; Information retrieval; Data mining; Artificial intelligence; Machine learning; Programming language","score_opus":0.01472741856136428,"score_gpt":0.2683685151964147,"score_spread":0.2536410966350504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953431343","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040590152,0.001613171,0.7480807,0.0010108444,0.00025638295,0.0014529246,0.088080265,0.094366305,0.024549233],"genre_scores_gemma":[0.13430981,0.0012873717,0.70207965,0.0004280361,0.00014871133,0.0015354531,0.14914142,0.0046237996,0.0064456896],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99614656,0.0005626896,0.0009378187,0.0006872453,0.0014207362,0.00024501915],"domain_scores_gemma":[0.99398583,0.0023165243,0.0007766295,0.0011997019,0.0014275046,0.00029391673],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022397542,0.0019017436,0.001097745,0.02374663,0.0015074156,0.003755121,0.0020090423,0.0017295453,0.010708933],"category_scores_gemma":[0.019065795,0.0008404555,0.0024061175,0.015208784,0.00075927144,0.011072604,0.0049165883,0.0011808363,0.006958304],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079839036,0.00033202957,0.023150519,0.0046747043,0.000508558,0.0016855503,0.007631649,0.011985802,0.028512836,0.08948096,0.22148205,0.60975695],"study_design_scores_gemma":[0.00022594245,0.00032872707,0.013788639,0.001310609,0.000352442,0.0023787115,0.005189854,0.26658922,0.029537069,0.17414957,0.5057692,0.00038008473],"about_ca_topic_score_codex":0.013084336,"about_ca_topic_score_gemma":0.014973485,"teacher_disagreement_score":0.02374663,"about_ca_system_score_codex":0.0018029965,"about_ca_system_score_gemma":0.0020863346,"threshold_uncertainty_score":0.035824955},"labels":[],"label_agreement":null},{"id":"W2953519948","doi":"10.1109/msr.2019.00066","title":"A Dataset of Non-Functional Bugs","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Python (programming language); Java; Programming language; Functional programming; Open source; Source code; Software engineering; Information retrieval; Software","score_opus":0.015193771002568659,"score_gpt":0.2579667185057456,"score_spread":0.24277294750317693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953519948","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09988728,0.002030349,0.0038666974,0.00052829186,0.00013342338,0.00026966535,0.88653725,0.0031137837,0.0036332596],"genre_scores_gemma":[0.035559017,0.00033819405,0.0056749443,0.00017717261,0.000031914205,0.00037400355,0.9562288,0.000180031,0.0014360019],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99705327,0.0003502927,0.00039639045,0.00070524134,0.0012380849,0.0002566853],"domain_scores_gemma":[0.99035734,0.0026552982,0.0018471401,0.0016195492,0.0025027294,0.0010179429],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014397646,0.0014866458,0.0009273174,0.0064917067,0.0008531115,0.0010204765,0.0018924943,0.0022194237,0.0035162922],"category_scores_gemma":[0.010923228,0.00045005637,0.0010749743,0.005883486,0.0005950876,0.0010324317,0.0017576463,0.0015184402,0.0053340164],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017284899,0.0015481439,0.23460244,0.006181053,0.0007181568,0.0028963878,0.0012053301,0.008640195,0.013486191,0.0028373045,0.62362015,0.10253614],"study_design_scores_gemma":[0.0005293844,0.0008110785,0.36246535,0.0008093315,0.00036469309,0.0036175926,0.0010701152,0.012554643,0.00896006,0.0029432045,0.6056486,0.00022586806],"about_ca_topic_score_codex":0.011920121,"about_ca_topic_score_gemma":0.025068155,"teacher_disagreement_score":0.011920121,"about_ca_system_score_codex":0.00092644093,"about_ca_system_score_gemma":0.0020726337,"threshold_uncertainty_score":0.023701489},"labels":[],"label_agreement":null},{"id":"W2953886423","doi":"10.1109/icpc.2019.00022","title":"Comparing Bug Replication in Regular and Micro Code Clones","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Replication (statistics); Cloning (programming); Copying; Computer science; Biology; Programming language; Java; Software bug; Code (set theory); Software maintenance; Software development; Software; Genetics; Gene; Virology","score_opus":0.023146019144085105,"score_gpt":0.26962647105587145,"score_spread":0.24648045191178636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953886423","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9964276,0.00027658933,0.0027454922,0.000014176878,0.000006851118,0.00003653255,0.00008500313,0.0000841058,0.00032372918],"genre_scores_gemma":[0.9947772,0.00009985647,0.0044103665,0.0000107684455,0.0000068359227,0.000036971287,0.00036512184,0.000022805561,0.00027004918],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9942211,0.001418362,0.0006237324,0.001572213,0.0018671134,0.00029736123],"domain_scores_gemma":[0.9266686,0.04878662,0.010571407,0.006223132,0.0064156563,0.0013346602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003952117,0.00031816968,0.00046907776,0.00400727,0.0004122322,0.0007701577,0.0005900605,0.00056668685,0.0005659102],"category_scores_gemma":[0.035637617,0.0002183449,0.00045835678,0.0020956513,0.0008173473,0.0016126649,0.0011246637,0.00044009794,0.00011322315],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001541038,0.00032536138,0.8043912,0.0005968849,0.00046983434,0.00079033454,0.0040650847,0.0071644397,0.025197119,0.0010575018,0.0006199309,0.15378128],"study_design_scores_gemma":[0.00006866705,0.0020550021,0.93679327,0.00006461989,0.000302756,0.001982394,0.0023210575,0.036431547,0.016377375,0.0013819741,0.0021494275,0.000071933355],"about_ca_topic_score_codex":0.0013389859,"about_ca_topic_score_gemma":0.0021874874,"teacher_disagreement_score":0.00400727,"about_ca_system_score_codex":0.0004957405,"about_ca_system_score_gemma":0.00038609543,"threshold_uncertainty_score":0.020901024},"labels":[],"label_agreement":null},{"id":"W2954059837","doi":"10.1109/msr.2019.00047","title":"How Often and What StackOverflow Posts Do Developers Reference in Their GitHub Projects?","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Maintainability; Code review; JavaScript; World Wide Web; Code (set theory); Code reuse; Source code; Software; Reuse; Software engineering; Static program analysis; Database; Software development; Programming language; Engineering","score_opus":0.024517222827480485,"score_gpt":0.24042222940018604,"score_spread":0.21590500657270556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954059837","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9869366,0.0006818719,0.0014674958,0.00037835052,0.000046416924,0.000029673363,0.0065164496,0.00060723803,0.0033357744],"genre_scores_gemma":[0.9629527,0.00092726335,0.0057336544,0.00020360756,0.00012166002,0.00012446044,0.022793466,0.0005583188,0.0065850215],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9987489,0.00027848574,0.00010674091,0.0003115738,0.00038405278,0.00017018736],"domain_scores_gemma":[0.9873055,0.0054694694,0.0042146244,0.00068169204,0.0016490247,0.00067965983],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0012509854,0.00036323673,0.0002819889,0.004766542,0.0005283837,0.0009594049,0.0003430391,0.0005928317,0.0015342042],"category_scores_gemma":[0.015419455,0.00019841865,0.00021778206,0.0034428583,0.0004150342,0.0019403471,0.0011078871,0.00040995947,0.0012326392],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035981752,0.000076729186,0.8449958,0.0008712015,0.00011779339,0.00078372966,0.014903074,0.00041281202,0.00875584,0.0009477289,0.028298786,0.0994767],"study_design_scores_gemma":[0.0000121881485,0.00008404283,0.95394456,0.0001810049,0.00006024053,0.0006640867,0.0061041377,0.0026620135,0.0026184558,0.0005988,0.033019032,0.000051476756],"about_ca_topic_score_codex":0.0057682153,"about_ca_topic_score_gemma":0.01817557,"teacher_disagreement_score":0.998749,"about_ca_system_score_codex":0.00042595586,"about_ca_system_score_gemma":0.00046489187,"threshold_uncertainty_score":0.011469245},"labels":[],"label_agreement":null},{"id":"W2954101277","doi":"10.1109/formalise.2019.00019","title":"A Vision for Helping Developers Use APIs by Leveraging Temporal Patterns","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Documentation; Leverage (statistics); Android (operating system); Coding (social sciences); World Wide Web; Human–computer interaction; Process (computing); Software engineering; Data science; Artificial intelligence; Programming language","score_opus":0.027274001645584016,"score_gpt":0.27595697800767843,"score_spread":0.2486829763620944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954101277","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01334889,0.0015967145,0.95598537,0.016513037,0.00020769419,0.00030886754,0.00023251153,0.00395766,0.007849342],"genre_scores_gemma":[0.06629846,0.0009674393,0.9271508,0.0011844349,0.00006758293,0.00028882406,0.00041548937,0.0002754063,0.0033516015],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9934362,0.0025672563,0.0006037409,0.0015185068,0.0015645705,0.0003097424],"domain_scores_gemma":[0.97674906,0.007868764,0.0029188783,0.005390911,0.0047107884,0.002361585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011084212,0.0020507202,0.0007771853,0.0042698984,0.0020253197,0.006267039,0.003445847,0.005225016,0.003613542],"category_scores_gemma":[0.025367053,0.0020712328,0.0019061791,0.0026760267,0.003825841,0.014682944,0.0043936837,0.0062445723,0.003142552],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038987814,0.0016847973,0.023341363,0.0019117398,0.00029551153,0.00066042104,0.013033572,0.013231233,0.030496372,0.16198495,0.030046146,0.72292405],"study_design_scores_gemma":[0.00027173883,0.0012748312,0.014190131,0.0013108401,0.00052634784,0.0030033502,0.00711711,0.18885884,0.01676291,0.483345,0.28270313,0.00063575833],"about_ca_topic_score_codex":0.009114283,"about_ca_topic_score_gemma":0.010547737,"teacher_disagreement_score":0.011084212,"about_ca_system_score_codex":0.0017419163,"about_ca_system_score_gemma":0.004649588,"threshold_uncertainty_score":0.05861962},"labels":[],"label_agreement":null},{"id":"W2954137348","doi":"10.1109/msr.2019.00084","title":"Scalable Software Merging Studies with MERGANSER","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Merge (version control); Scalability; Scripting language; Python (programming language); SQL; Software; Merge algorithm; Database; Software engineering; Data mining; Information retrieval; Programming language","score_opus":0.017473481324697844,"score_gpt":0.2654499640269312,"score_spread":0.24797648270223338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954137348","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09152579,0.0011335275,0.70559967,0.0017391117,0.0002775197,0.0018265947,0.011639679,0.17133194,0.014926097],"genre_scores_gemma":[0.18988185,0.00046222084,0.783642,0.0002268049,0.00007383388,0.0012816732,0.013431448,0.008084679,0.0029154504],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9914971,0.0027350679,0.0011377216,0.0017729384,0.0026147077,0.00024241884],"domain_scores_gemma":[0.95981926,0.022151053,0.002890898,0.011033126,0.003378365,0.0007273569],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013700479,0.001532118,0.0012356867,0.008919726,0.0014155565,0.0046513756,0.0034447224,0.0010938597,0.012189538],"category_scores_gemma":[0.049101524,0.0013056266,0.0017021537,0.004689573,0.0011421085,0.0075593838,0.008055645,0.0017196988,0.0025103828],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00095063075,0.0007576437,0.049071554,0.0032635613,0.0010937911,0.0017156469,0.009750322,0.04774055,0.02139479,0.06800036,0.069726974,0.72653425],"study_design_scores_gemma":[0.0007711646,0.00061664963,0.025745831,0.0012283089,0.00051002816,0.001431209,0.0040496658,0.5051361,0.049077135,0.2051549,0.20574312,0.0005359663],"about_ca_topic_score_codex":0.0035275044,"about_ca_topic_score_gemma":0.006895994,"teacher_disagreement_score":0.013700479,"about_ca_system_score_codex":0.0012526034,"about_ca_system_score_gemma":0.003179079,"threshold_uncertainty_score":0.07245594},"labels":[],"label_agreement":null},{"id":"W2954274464","doi":"10.1109/icse.2019.00022","title":"Natural Software Revisited","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Codebase; Programming language; Java; Scripting language; Source code; Code (set theory); Abstract syntax; Punctuation; Artificial intelligence; Semantics (computer science)","score_opus":0.008383827925464686,"score_gpt":0.2448931584466986,"score_spread":0.2365093305212339,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954274464","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03516978,0.021102263,0.19376478,0.049802545,0.0021341715,0.00016725555,0.0020326958,0.0012760375,0.6945505],"genre_scores_gemma":[0.7949966,0.007885507,0.06914511,0.013998178,0.0014590457,0.00045602611,0.0027222424,0.0012914801,0.10804582],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9946109,0.0022034994,0.0003189384,0.0013041926,0.0013715834,0.00019084426],"domain_scores_gemma":[0.9903312,0.0052884487,0.00073973753,0.0018279841,0.0015188362,0.00029386475],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025167884,0.0006662189,0.00056714297,0.0039137714,0.0029366547,0.0053381063,0.00120562,0.0021149595,0.016338212],"category_scores_gemma":[0.014426399,0.0003461848,0.00083167007,0.00316515,0.014846579,0.009519689,0.0032024009,0.0029265443,0.0033497438],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012576816,0.000007814852,0.0006535248,0.00016451678,0.000008574492,0.00020793102,0.0026208113,0.00031254505,0.0002058367,0.9572115,0.012992549,0.025601912],"study_design_scores_gemma":[0.00000697089,0.000015145562,0.0008446313,0.00020632675,0.0000069702983,0.0009181589,0.001553316,0.0014962554,0.00022896494,0.4880786,0.5066261,0.000018575198],"about_ca_topic_score_codex":0.00768324,"about_ca_topic_score_gemma":0.006461463,"teacher_disagreement_score":0.016338212,"about_ca_system_score_codex":0.005041236,"about_ca_system_score_gemma":0.0025382023,"threshold_uncertainty_score":0.054656744},"labels":[],"label_agreement":null},{"id":"W2954476116","doi":"10.1109/msr.2019.00046","title":"Can Duplicate Questions on Stack Overflow Benefit the Software Development Community?","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Heuristics; Reputation; Information retrieval; Similarity (geometry); Stack (abstract data type); World Wide Web; Data science; Data mining; Artificial intelligence; Programming language","score_opus":0.02334161124921441,"score_gpt":0.2581692911669891,"score_spread":0.2348276799177747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954476116","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9687413,0.0011356183,0.01321449,0.0041320715,0.00017718003,0.00019229628,0.00081546936,0.0009184247,0.010673275],"genre_scores_gemma":[0.9790589,0.00042497437,0.012334731,0.001325405,0.00025509443,0.00009620541,0.0014566495,0.0002862625,0.0047617825],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97861326,0.010580679,0.0014183219,0.0022565268,0.0058418885,0.0012893834],"domain_scores_gemma":[0.851664,0.09782714,0.018400326,0.012794448,0.013864786,0.0054492503],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.018721145,0.00061875116,0.00092450355,0.0038263826,0.0030180342,0.0042562517,0.0014663341,0.0028442931,0.0042161485],"category_scores_gemma":[0.16799183,0.0005477974,0.0006895422,0.0028438536,0.002056108,0.01320161,0.005910259,0.0017066575,0.0014633934],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013476658,0.00068298005,0.50842655,0.0016489985,0.0002654335,0.0033409246,0.09697936,0.0023258089,0.01269505,0.010050863,0.030625029,0.3316114],"study_design_scores_gemma":[0.00035046248,0.0025208874,0.44219065,0.0013108804,0.0005573479,0.010573456,0.14438854,0.040409897,0.022217242,0.052785628,0.28202665,0.0006683456],"about_ca_topic_score_codex":0.0036808478,"about_ca_topic_score_gemma":0.0054458105,"teacher_disagreement_score":0.98127884,"about_ca_system_score_codex":0.0014858369,"about_ca_system_score_gemma":0.002371155,"threshold_uncertainty_score":0.09900808},"labels":[],"label_agreement":null},{"id":"W2954536659","doi":"10.1109/icse-companion.2019.00128","title":"Constructural Software Documentation","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Documentation; Unit testing; Header; Software documentation; Redundancy (engineering); Software; Software engineering; ENCODE; Programming language; Software development; Software construction; Operating system","score_opus":0.007624648879982553,"score_gpt":0.2550371852956053,"score_spread":0.24741253641562277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954536659","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029130118,0.00039279295,0.9042277,0.0006435654,0.000101082645,0.00020736679,0.0036553682,0.030831182,0.030810837],"genre_scores_gemma":[0.37539914,0.00072252605,0.5737068,0.00052459096,0.00012406573,0.00042716513,0.013274464,0.011503013,0.024318358],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9959996,0.0010908769,0.0005658747,0.00053494965,0.0016236125,0.00018504371],"domain_scores_gemma":[0.9708282,0.008420736,0.0023671284,0.013246937,0.0047356365,0.00040128987],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038537194,0.00093759724,0.0004002328,0.003004906,0.00081540446,0.004084794,0.0016343489,0.001253501,0.008786095],"category_scores_gemma":[0.024638478,0.00086997024,0.00052525406,0.00263909,0.0015210878,0.004553011,0.0032606234,0.0019934722,0.0039494284],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035125294,0.00027051117,0.014664308,0.0014018185,0.00008398407,0.00071556686,0.0052162516,0.0056352415,0.027869988,0.2757247,0.03771197,0.63035446],"study_design_scores_gemma":[0.00010015332,0.00019489875,0.013192084,0.0012942454,0.00012371023,0.0017014643,0.0006505451,0.037299514,0.07352941,0.14736642,0.72436863,0.0001788838],"about_ca_topic_score_codex":0.0016992949,"about_ca_topic_score_gemma":0.0025536607,"teacher_disagreement_score":0.008786095,"about_ca_system_score_codex":0.00095545605,"about_ca_system_score_gemma":0.0020379333,"threshold_uncertainty_score":0.029392362},"labels":[],"label_agreement":null},{"id":"W2954574984","doi":"10.1109/msr.2019.00052","title":"What do Developers Know About Machine Learning: A Study of ML Discussions on StackOverflow","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Latent Dirichlet allocation; Machine learning; Artificial intelligence; Topic model; Software; Focus (optics); Data science","score_opus":0.013979527115837276,"score_gpt":0.27507822055860637,"score_spread":0.2610986934427691,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954574984","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9931977,0.00019315208,0.00045122762,0.0019586345,0.000020434114,0.000024125677,0.00013115785,0.000040330186,0.003983115],"genre_scores_gemma":[0.99549556,0.00024909264,0.0004006887,0.0008756818,0.000065463864,0.000053304746,0.00030668708,0.000072308576,0.002481161],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9959468,0.0024372495,0.00021107105,0.0003668846,0.0006050659,0.00043289267],"domain_scores_gemma":[0.9246711,0.053129792,0.010846953,0.0016820506,0.0052394974,0.004430582],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0062333737,0.00030295146,0.00037037476,0.0026497873,0.0038293628,0.0036317173,0.0007489638,0.001840631,0.0032546069],"category_scores_gemma":[0.05510268,0.0004243058,0.00020857164,0.0019564908,0.001744448,0.008890771,0.0028170845,0.0027156312,0.0011285901],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026460132,0.00053112727,0.19010499,0.00017965738,0.000025115713,0.0006370076,0.76222205,0.00012901347,0.0020981464,0.0014796498,0.006326404,0.036002163],"study_design_scores_gemma":[0.000059660346,0.00047635534,0.28613457,0.00027559645,0.000028624909,0.00044197042,0.65255934,0.0023277136,0.0012684003,0.0015955693,0.054713033,0.000119149154],"about_ca_topic_score_codex":0.0074117896,"about_ca_topic_score_gemma":0.010600039,"teacher_disagreement_score":0.9937666,"about_ca_system_score_codex":0.002038168,"about_ca_system_score_gemma":0.0010766678,"threshold_uncertainty_score":0.03296566},"labels":[],"label_agreement":null},{"id":"W2954709499","doi":"10.1109/msr.2019.00053","title":"Investigating Next Steps in Static API-Misuse Detection","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Precision and recall; Application programming interface; Ranking (information retrieval); Recall; Software; Graph; Data mining; Information retrieval; Programming language; Theoretical computer science","score_opus":0.026785942484286743,"score_gpt":0.2681808507987636,"score_spread":0.24139490831447685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954709499","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23262854,0.007896833,0.71621794,0.01360615,0.00046136137,0.00088031753,0.0009994378,0.018657975,0.008651393],"genre_scores_gemma":[0.44339982,0.0015032676,0.54706943,0.0014694531,0.00013095431,0.00031608136,0.0025563885,0.0007619374,0.0027925493],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98663634,0.004096068,0.0009018396,0.002691932,0.0042848014,0.0013890405],"domain_scores_gemma":[0.9448518,0.03018507,0.002808871,0.0072752144,0.013186017,0.0016929783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0150259435,0.0027253001,0.002340571,0.009030712,0.0023748474,0.00804248,0.0041153547,0.0034871257,0.0025026368],"category_scores_gemma":[0.052623186,0.0016892256,0.001929628,0.0043441695,0.002067058,0.016580252,0.0035437813,0.0044677057,0.002756091],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004757895,0.0016068621,0.12014954,0.00128346,0.0004148456,0.0007373109,0.0016279607,0.021900529,0.029483257,0.022817202,0.021786205,0.777717],"study_design_scores_gemma":[0.00013564179,0.0010627449,0.025074348,0.0007487249,0.0005319177,0.001772655,0.0037581977,0.7596339,0.06622788,0.08816816,0.052550543,0.00033531574],"about_ca_topic_score_codex":0.012173759,"about_ca_topic_score_gemma":0.016094038,"teacher_disagreement_score":0.0150259435,"about_ca_system_score_codex":0.0018702167,"about_ca_system_score_gemma":0.005103117,"threshold_uncertainty_score":0.07946575},"labels":[],"label_agreement":null},{"id":"W2954741184","doi":"10.1109/msr.2019.00080","title":"Predicting Co-Changes between Functionality Specifications and Source Code in Behavior Driven Development","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Traceability; Feature (linguistics); Source code; Documentation; Codebase; Programming language; Feature model; Software; Software development; Software engineering; Database","score_opus":0.0741375878839423,"score_gpt":0.2966963779232075,"score_spread":0.2225587900392652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954741184","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9544373,0.00037549806,0.042333845,0.00017894946,0.000021370652,0.0001175582,0.00091473217,0.0011250636,0.00049563614],"genre_scores_gemma":[0.96150124,0.00010832563,0.034546506,0.000049403287,0.000010202405,0.00012777909,0.0032333273,0.00007319948,0.00034989623],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970969,0.00093133276,0.0002000711,0.00079669966,0.0007710255,0.0002039514],"domain_scores_gemma":[0.9592884,0.030327715,0.0048880884,0.0013988385,0.003480749,0.0006162242],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039140037,0.00088183087,0.00049706566,0.004136395,0.0004054182,0.00078617147,0.0007442481,0.0009426694,0.00032171555],"category_scores_gemma":[0.026516754,0.0004101248,0.0009208048,0.0021880579,0.000383886,0.0012784865,0.0005890701,0.0009023976,0.0002761587],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005781978,0.0008268119,0.6408352,0.00034894922,0.00029561913,0.0008451226,0.000951211,0.16586256,0.008273191,0.0006806847,0.0024396002,0.17806287],"study_design_scores_gemma":[0.00002579575,0.000293412,0.10859673,0.00004426345,0.000062153325,0.00036499696,0.00018996042,0.8821463,0.0055847806,0.0014015957,0.0012453157,0.000044681885],"about_ca_topic_score_codex":0.013164279,"about_ca_topic_score_gemma":0.016542135,"teacher_disagreement_score":0.013164279,"about_ca_system_score_codex":0.0010191294,"about_ca_system_score_gemma":0.0010075624,"threshold_uncertainty_score":0.02617532},"labels":[],"label_agreement":null},{"id":"W2954905606","doi":"10.1109/icse-companion.2019.00125","title":"Towards Visualizing Large Scale Evolving Clones","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code refactoring; Computer science; clone (Java method); Software visualization; Software; Software engineering; Visualization; Software system; Software maintenance; Software development; Software evolution; Zoom; Data science; Software framework; Code (set theory); Software construction; Programming language; Data mining; Engineering","score_opus":0.01383199245937057,"score_gpt":0.290548167654324,"score_spread":0.27671617519495345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954905606","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054696567,0.0021107022,0.92274463,0.0022648135,0.00012552053,0.0001506867,0.00087004714,0.0123462025,0.00469072],"genre_scores_gemma":[0.26504147,0.0018362504,0.7282133,0.00033128168,0.00010823206,0.00017959005,0.0012009268,0.0012499031,0.0018390573],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99945825,0.00017052569,0.000027952663,0.00008770864,0.0002207539,0.000034786095],"domain_scores_gemma":[0.99554825,0.0018752242,0.0004978671,0.00048219325,0.0012141619,0.0003822248],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018668808,0.0010481956,0.00043938763,0.0058786105,0.0005925854,0.0043006763,0.0011695002,0.0013448476,0.0025330046],"category_scores_gemma":[0.01016746,0.00052011333,0.0006890915,0.0026649323,0.0009281942,0.0050571645,0.0033715118,0.0017562547,0.00056203536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062305765,0.0002320188,0.018953672,0.0014737785,0.00020599038,0.0014369088,0.023876658,0.059463393,0.06914673,0.11588858,0.039901547,0.6687976],"study_design_scores_gemma":[0.00013651271,0.00017370605,0.012048079,0.00075451663,0.00016601298,0.0018182384,0.006324124,0.59525436,0.040791985,0.19920757,0.14308324,0.00024165228],"about_ca_topic_score_codex":0.003190267,"about_ca_topic_score_gemma":0.0030516535,"teacher_disagreement_score":0.0058786105,"about_ca_system_score_codex":0.0007340942,"about_ca_system_score_gemma":0.000761176,"threshold_uncertainty_score":0.009873092},"labels":[],"label_agreement":null},{"id":"W2955136383","doi":"10.1007/s11219-019-09456-3","title":"Pieces of contextual information suitable for predicting co-changes? An empirical study","year":2019,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Commit; Computer science; Artifact (error); Metadata; Baseline (sea); Software; Software engineering; Set (abstract data type); Software development; Data science; Data mining; Information retrieval; Artificial intelligence; World Wide Web; Database; Programming language","score_opus":0.06126340530962837,"score_gpt":0.39334973887793756,"score_spread":0.3320863335683092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2955136383","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99734586,0.0002058521,0.00096880493,0.000044616136,0.000008744604,0.000051914656,0.0004098354,0.000024045888,0.0009401875],"genre_scores_gemma":[0.9981856,0.00007318009,0.0010772878,0.000011435364,0.000009976766,0.000028656475,0.00048487366,0.000009651314,0.0001193395],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9963742,0.0016054336,0.00035612207,0.00056974805,0.0008587647,0.00023580404],"domain_scores_gemma":[0.90701705,0.06651692,0.01194282,0.006101497,0.0062066256,0.0022151412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002956989,0.00038085107,0.0006928404,0.0042536883,0.0008013612,0.0016621368,0.0007876168,0.0010062845,0.0018849694],"category_scores_gemma":[0.044158652,0.00027131883,0.00047818647,0.0071565616,0.00057001435,0.0033078745,0.001344731,0.00082170335,0.0004357965],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010697294,0.00094444223,0.9487245,0.00020397958,0.00018057134,0.000347991,0.0027925754,0.0011515019,0.0015438335,0.00028948794,0.00041494827,0.042336393],"study_design_scores_gemma":[0.00003254128,0.0006482215,0.97727215,0.00007514185,0.0003097505,0.00031081963,0.0051734135,0.0119781615,0.0015292319,0.0006279981,0.0019988806,0.00004367818],"about_ca_topic_score_codex":0.0050307666,"about_ca_topic_score_gemma":0.007489387,"teacher_disagreement_score":0.0050307666,"about_ca_system_score_codex":0.00061167777,"about_ca_system_score_gemma":0.00065696775,"threshold_uncertainty_score":0.015638232},"labels":[],"label_agreement":null},{"id":"W2955343015","doi":"10.1109/icpc.2019.00030","title":"Visualizing Sequences of Debugging Sessions using Swarm Debugging","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Debugging; Computer science; Algorithmic program debugging; Session (web analytics); Visualization; Microsoft Visual Studio; Software engineering; Programming language; Background debug mode interface; Program comprehension; Software; Software bug; Human–computer interaction; Software system; World Wide Web; Artificial intelligence","score_opus":0.034793687281973046,"score_gpt":0.3293267699547837,"score_spread":0.29453308267281064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2955343015","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08915839,0.0006490424,0.88445854,0.00032715494,0.00015724964,0.00024820582,0.0019692727,0.017044822,0.0059873024],"genre_scores_gemma":[0.42648324,0.00062817504,0.56572604,0.000079457815,0.0000622079,0.00025352504,0.0023884678,0.0016430006,0.002735881],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994186,0.00020323083,0.000049524886,0.00012715143,0.00014200603,0.000059495982],"domain_scores_gemma":[0.99754024,0.0012887995,0.00025121265,0.00033161492,0.0003691284,0.00021900814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012121542,0.0010102397,0.00043719917,0.0033811794,0.00045600982,0.0014202981,0.0007094659,0.0007574023,0.004574372],"category_scores_gemma":[0.0034475666,0.00034466173,0.0005822637,0.0016419033,0.0003454162,0.0016968343,0.0014078586,0.00073880085,0.0006158008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018479929,0.000498467,0.032395203,0.0015975485,0.0002524518,0.0016723429,0.015869448,0.07794029,0.08696575,0.02923388,0.027575089,0.72415155],"study_design_scores_gemma":[0.00047741758,0.0010688737,0.03591171,0.00069822,0.00028683763,0.0019222662,0.0044302824,0.650998,0.09214808,0.044878915,0.16676079,0.00041865782],"about_ca_topic_score_codex":0.0022430462,"about_ca_topic_score_gemma":0.002603294,"teacher_disagreement_score":0.004574372,"about_ca_system_score_codex":0.00030496693,"about_ca_system_score_gemma":0.0008159237,"threshold_uncertainty_score":0.015302777},"labels":[],"label_agreement":null},{"id":"W2955628496","doi":"10.1109/icse-companion.2019.00021","title":"Publish or Perish: Questioning the Impact of Our Research on the Software Developer","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Publication; Context (archaeology); Software peer review; Software; Publish or perish; Social software engineering; Software walkthrough; Publishing; Software engineering; Productivity; Software development; World Wide Web; Data science; Software construction; Business","score_opus":0.1066703820197611,"score_gpt":0.39779859058050804,"score_spread":0.29112820856074695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2955628496","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016351515,0.018852727,0.016032817,0.9075543,0.008596157,0.000103794206,0.000061523606,0.00024280159,0.032204393],"genre_scores_gemma":[0.6956367,0.03455508,0.035406064,0.19592056,0.019761099,0.00090223225,0.0001518364,0.0023189357,0.01534745],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.65299684,0.25179395,0.010686741,0.017004406,0.06158919,0.005928877],"domain_scores_gemma":[0.18344761,0.69288576,0.016203834,0.041862294,0.05284409,0.012756398],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.3628474,0.001236743,0.0024609612,0.008619915,0.024247965,0.066013105,0.00801066,0.018757507,0.007138759],"category_scores_gemma":[0.62444156,0.0015731893,0.001472862,0.009107085,0.07266196,0.11652263,0.024572583,0.027117819,0.0031902744],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020389588,0.00023715512,0.008990991,0.0021194066,0.00025394608,0.00044669028,0.15254591,0.0005116648,0.0009898078,0.5993266,0.08983284,0.14454108],"study_design_scores_gemma":[0.00011495198,0.0003893503,0.0029724743,0.0044424045,0.00019396583,0.00060101657,0.16324021,0.0012905578,0.0022235992,0.5158421,0.30847666,0.00021268087],"about_ca_topic_score_codex":0.0029852276,"about_ca_topic_score_gemma":0.0029881871,"teacher_disagreement_score":0.6371526,"about_ca_system_score_codex":0.012791066,"about_ca_system_score_gemma":0.023152689,"threshold_uncertainty_score":0.7857226},"labels":[],"label_agreement":null},{"id":"W2955818766","doi":"10.1109/icsme.2019.00072","title":"MigrationMiner: An Automated Detection Tool of Third-Party Java Library Migration at the Method Level","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Java; Computer science; Documentation; Benchmark (surveying); Source code; Code (set theory); Process (computing); Open source; Programming language; World Wide Web; Software; Geography; Cartography","score_opus":0.039210357538860735,"score_gpt":0.3142249769733309,"score_spread":0.27501461943447014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2955818766","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08388479,0.0023913793,0.4373089,0.00067707326,0.00035155992,0.0004206913,0.0059320647,0.46481237,0.0042212163],"genre_scores_gemma":[0.34532195,0.0012047479,0.6031328,0.0010549919,0.00014888181,0.00069436367,0.016998796,0.021311333,0.010132119],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99564266,0.00063783035,0.0004095751,0.0010454673,0.002010996,0.00025350845],"domain_scores_gemma":[0.9892852,0.0042549735,0.0024173888,0.002127504,0.0016342619,0.00028068593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002696815,0.0019481194,0.0009913354,0.0050817435,0.0007188767,0.0016460803,0.0023235162,0.0019460121,0.0023070134],"category_scores_gemma":[0.015395679,0.0011864924,0.0011359456,0.0014368041,0.0005932011,0.0026792155,0.0023573125,0.0017165189,0.0029199105],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008528361,0.00048718075,0.07247432,0.0020830184,0.00036118052,0.0019913507,0.0021162408,0.008963574,0.07055277,0.005000836,0.13362685,0.70148987],"study_design_scores_gemma":[0.0002615574,0.00061746925,0.057581823,0.0010063446,0.00029540138,0.0065047396,0.0008207706,0.39389712,0.27170733,0.011917896,0.25488645,0.0005031032],"about_ca_topic_score_codex":0.0028695902,"about_ca_topic_score_gemma":0.004025775,"teacher_disagreement_score":0.0050817435,"about_ca_system_score_codex":0.0006504773,"about_ca_system_score_gemma":0.0017465757,"threshold_uncertainty_score":0.014262259},"labels":[],"label_agreement":null},{"id":"W2955851367","doi":"10.1109/msr.2019.00074","title":"Can Issues Reported at Stack Overflow Questions be Reproduced? An Exploratory Study","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Code (set theory); Code review; Compiler; Exploratory research; Programming language; Stack (abstract data type); Source code; Sample (material); Software; Data science; Software development; Static program analysis","score_opus":0.05047758650881961,"score_gpt":0.3264721570450574,"score_spread":0.27599457053623777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2955851367","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99587244,0.000121075704,0.002617617,0.0001839425,0.000009302289,0.00030168993,0.00014373542,0.000033381664,0.00071684615],"genre_scores_gemma":[0.9954829,0.00013344947,0.002940288,0.00017777871,0.000027711658,0.00061275635,0.00024202933,0.00004435064,0.0003386848],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9234185,0.05273058,0.0071132537,0.004613916,0.010327937,0.0017958428],"domain_scores_gemma":[0.38539875,0.51332873,0.058259573,0.018457048,0.022676658,0.001879165],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06842841,0.00071176334,0.000660761,0.004729498,0.0014130824,0.0027013305,0.0016210675,0.0018338688,0.0014354809],"category_scores_gemma":[0.38103747,0.0007446841,0.0008394427,0.002634826,0.0027525353,0.0036890707,0.003408573,0.0016780979,0.00047292744],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006714481,0.0013978762,0.54464257,0.0011247203,0.00028834096,0.002238647,0.39867917,0.00045812503,0.005941588,0.0009595412,0.0013143113,0.04228367],"study_design_scores_gemma":[0.00011940991,0.005244932,0.67080325,0.0012644494,0.00030990303,0.004808476,0.27740765,0.0043378677,0.013193107,0.0020211511,0.020198185,0.00029158455],"about_ca_topic_score_codex":0.00056975824,"about_ca_topic_score_gemma":0.00077559164,"teacher_disagreement_score":0.9315716,"about_ca_system_score_codex":0.0011426498,"about_ca_system_score_gemma":0.0010536746,"threshold_uncertainty_score":0.36188835},"labels":[],"label_agreement":null},{"id":"W2956081399","doi":"10.1109/icse-companion.2019.00088","title":"Supporting Code Search with Context-Aware, Analytics-Driven, Effective Query Reformulation","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Analytics; Context (archaeology); Code (set theory); Information retrieval; Web search query; Selection (genetic algorithm); Data science; World Wide Web; Search engine; Programming language; Machine learning","score_opus":0.014327123551129453,"score_gpt":0.2973555347354912,"score_spread":0.28302841118436173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2956081399","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13626732,0.0020296427,0.8268297,0.0038213602,0.00017301933,0.002224449,0.0020425885,0.021062678,0.0055493005],"genre_scores_gemma":[0.37988353,0.00061963225,0.6123375,0.0006045559,0.00013341353,0.0006687202,0.0033347118,0.0007951313,0.0016228545],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98806155,0.005400923,0.0011284796,0.0017770346,0.003122626,0.0005094084],"domain_scores_gemma":[0.96301454,0.022011353,0.0024605566,0.0059255525,0.00567244,0.0009156164],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00902031,0.0018089182,0.0019327898,0.006788749,0.0012533194,0.0044991146,0.002448111,0.0015351468,0.0024579265],"category_scores_gemma":[0.062386326,0.00071694975,0.0012915917,0.0038307672,0.0011550571,0.008481774,0.005009034,0.0020983107,0.0017872602],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001475077,0.0013951105,0.013770389,0.0030445613,0.00027781873,0.0008729263,0.013057888,0.030162128,0.101174034,0.01668904,0.029288732,0.7887922],"study_design_scores_gemma":[0.0003737397,0.0008302939,0.007024886,0.0005032033,0.00037400282,0.0008515558,0.008485996,0.78736746,0.08184367,0.0488157,0.063190214,0.00033931664],"about_ca_topic_score_codex":0.0067460123,"about_ca_topic_score_gemma":0.009088351,"teacher_disagreement_score":0.00902031,"about_ca_system_score_codex":0.0014001075,"about_ca_system_score_gemma":0.004184421,"threshold_uncertainty_score":0.047704577},"labels":[],"label_agreement":null},{"id":"W2956228533","doi":"10.18280/isi.240112","title":"An Analysis of Maintainability Index Influencing Metrics and Their Behavior on Similar Open Source Gaming Application Developed in C, C++ and, JAVA","year":2019,"lang":"en","type":"article","venue":"Ingénierie des systèmes d information","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Maintainability; Java; Open source; Index (typography); Computer science; Operating system; World Wide Web; Statistics; Software engineering; Software; Mathematics","score_opus":0.014117080524436651,"score_gpt":0.2692895046199603,"score_spread":0.25517242409552365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2956228533","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9972345,0.00008638167,0.0014368511,0.00003408395,0.0000056510316,0.000022999207,0.00030869982,0.00011425118,0.00075659336],"genre_scores_gemma":[0.99610114,0.000044069882,0.002361244,0.000007762818,0.0000064179876,0.00004188128,0.0008396163,0.000020488887,0.0005773091],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983375,0.00030216345,0.00017181096,0.00034950496,0.00071446,0.00012447636],"domain_scores_gemma":[0.9809209,0.007889154,0.0038831425,0.001112227,0.0053978534,0.0007967535],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014903771,0.0003892779,0.0002536969,0.0028256974,0.0003430471,0.00062554906,0.00024199531,0.00028942953,0.0004053156],"category_scores_gemma":[0.01400524,0.00014138667,0.00031651204,0.0021733844,0.00025698243,0.0006805498,0.00033465007,0.00048824464,0.00013497578],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016385966,0.00026592915,0.9111038,0.00012204501,0.00008693646,0.0002565626,0.0036832022,0.001520255,0.011073711,0.0003221916,0.0011318054,0.07026959],"study_design_scores_gemma":[0.000002684523,0.00029132777,0.9859446,0.000013674515,0.000034279477,0.0002349923,0.00086010486,0.0073157283,0.00386592,0.00015608544,0.0012535284,0.000027171895],"about_ca_topic_score_codex":0.0035366742,"about_ca_topic_score_gemma":0.005273867,"teacher_disagreement_score":0.0035366742,"about_ca_system_score_codex":0.0004105733,"about_ca_system_score_gemma":0.0003101019,"threshold_uncertainty_score":0.007881999},"labels":[],"label_agreement":null},{"id":"W2956548897","doi":"10.48550/arxiv.1907.04908","title":"Executability of Python Snippets in Stack Overflow","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Python (programming language); Executable; Computer science; Scalability; Documentation; Programming language; World Wide Web; Plug-in; Code (set theory); Information retrieval; Database","score_opus":0.0630678027969123,"score_gpt":0.20483467774349454,"score_spread":0.14176687494658224,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2956548897","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6962656,0.0008732214,0.05960435,0.0015056388,0.0004758737,0.0005459328,0.028830957,0.20069501,0.011203444],"genre_scores_gemma":[0.76421154,0.0004372249,0.11440536,0.0008377267,0.00018734459,0.00070877833,0.059302613,0.0470097,0.012899734],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99595946,0.000784151,0.0004327331,0.00070841296,0.0017496599,0.00036557062],"domain_scores_gemma":[0.9643026,0.02664272,0.002678923,0.0026911667,0.0030521946,0.00063252466],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004149208,0.0015627012,0.0007015604,0.003344782,0.0012691191,0.001511335,0.0011604163,0.001266182,0.0071232864],"category_scores_gemma":[0.03309327,0.000602384,0.0011606965,0.0016015338,0.0013095713,0.0029897215,0.0017501276,0.0013197478,0.00261796],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0077685863,0.0012080034,0.16046132,0.0052789617,0.0006699597,0.02181106,0.013966177,0.039292835,0.1001328,0.01612088,0.20293312,0.43035638],"study_design_scores_gemma":[0.00045742033,0.00086045853,0.15782894,0.0018239461,0.00041815033,0.0059959567,0.0042400775,0.33729368,0.2860631,0.03140921,0.17285937,0.0007497098],"about_ca_topic_score_codex":0.0037470418,"about_ca_topic_score_gemma":0.007963397,"teacher_disagreement_score":0.9958508,"about_ca_system_score_codex":0.00086506316,"about_ca_system_score_gemma":0.0016407102,"threshold_uncertainty_score":0.023829758},"labels":[],"label_agreement":null},{"id":"W2959431065","doi":"10.1109/scam.2019.00010","title":"A Study on the Effects of Exception Usage in Open-Source C++ Systems","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Java; Programming language; Modular design; Source code; Control flow; Python (programming language); Open source; Call graph; Exception handling; Control flow graph; Theoretical computer science; Empirical research; Exploratory research; Software engineering; Software","score_opus":0.022151142329705623,"score_gpt":0.2831149180961701,"score_spread":0.2609637757664645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2959431065","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9981857,0.00010414826,0.0007551407,0.000056411736,0.0000051680654,0.00003598344,0.00006288481,0.0001287899,0.00066575024],"genre_scores_gemma":[0.9969388,0.00008135937,0.0022025988,0.00004394605,0.000008583889,0.000036418936,0.00024085375,0.00010523044,0.00034228223],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.984009,0.0080107115,0.0011776881,0.0020952215,0.0039683194,0.0007390061],"domain_scores_gemma":[0.55364114,0.3720622,0.033427,0.020151313,0.017725809,0.002992557],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00735206,0.00071612425,0.00039841034,0.0020299389,0.001452051,0.0018076313,0.0014077739,0.0010299585,0.0012078442],"category_scores_gemma":[0.12510808,0.00056811655,0.00066087313,0.0026162388,0.0020655538,0.003036348,0.0016747351,0.0019877742,0.000365696],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025171738,0.0037299665,0.752096,0.0012905899,0.00052667264,0.0018429145,0.017615909,0.028628524,0.028922817,0.0013146709,0.0032760035,0.15823886],"study_design_scores_gemma":[0.00011235101,0.003612668,0.9294706,0.0001416041,0.00027272213,0.0010762798,0.007565468,0.036900327,0.015067498,0.0008378823,0.0047570914,0.00018555533],"about_ca_topic_score_codex":0.0068470617,"about_ca_topic_score_gemma":0.0078267865,"teacher_disagreement_score":0.00735206,"about_ca_system_score_codex":0.0014323907,"about_ca_system_score_gemma":0.000760339,"threshold_uncertainty_score":0.03888184},"labels":[],"label_agreement":null},{"id":"W2959891429","doi":"10.1109/icsme.2019.00048","title":"Syntax and Stack Overflow: A Methodology for Extracting a Corpus of Syntax Errors and Fixes","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Syntax error; Computer science; Python (programming language); Syntax; Abstract syntax tree; Parsing; Programming language; Natural language processing; Artificial intelligence; Abstract syntax; Source code","score_opus":0.09427978240340419,"score_gpt":0.3505228612023415,"score_spread":0.2562430787989373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2959891429","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.098185346,0.0012197004,0.3851356,0.0015676529,0.00066485227,0.0055757426,0.43467814,0.05557961,0.017393379],"genre_scores_gemma":[0.051864583,0.00046754206,0.5641948,0.00043800793,0.00014688837,0.007767276,0.36203846,0.0065866313,0.006495789],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98464173,0.0028267074,0.0038794759,0.0035505104,0.004523292,0.00057822495],"domain_scores_gemma":[0.9460947,0.021562845,0.007728657,0.011965585,0.011708474,0.0009397239],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009192422,0.0019223867,0.0010905104,0.020679794,0.002310793,0.0032456133,0.002288743,0.0021710848,0.007872728],"category_scores_gemma":[0.054859687,0.0013221262,0.0019167581,0.01595701,0.002324949,0.005088875,0.0061788578,0.0031565393,0.009604675],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007151351,0.00069325697,0.056039043,0.009544227,0.0005012681,0.002325414,0.016737508,0.005803806,0.04514198,0.026859108,0.3234942,0.51214504],"study_design_scores_gemma":[0.00025806073,0.00033848896,0.11417284,0.0012337811,0.00027235184,0.0020020248,0.00526405,0.03247215,0.05457691,0.032618947,0.75620157,0.00058884855],"about_ca_topic_score_codex":0.008034229,"about_ca_topic_score_gemma":0.0150850145,"teacher_disagreement_score":0.020679794,"about_ca_system_score_codex":0.0016292592,"about_ca_system_score_gemma":0.00789493,"threshold_uncertainty_score":0.04861474},"labels":[],"label_agreement":null},{"id":"W2960587195","doi":"10.1109/esem.2019.8870177","title":"On the Impact of Refactoring on the Relationship between Quality Attributes and Design Metrics","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Quality (philosophy); Software quality; Software engineering; Software; Set (abstract data type); Java; Software maintenance; Variety (cybernetics); Software metric; Software system; Software development; Data mining; Artificial intelligence; Programming language","score_opus":0.3496185331524583,"score_gpt":0.40854230430446214,"score_spread":0.05892377115200387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2960587195","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97562146,0.0033761538,0.014451451,0.00062184467,0.000046727902,0.00004490441,0.0022482795,0.00026067777,0.0033284281],"genre_scores_gemma":[0.99385864,0.00033752166,0.0039320854,0.000045768284,0.000033169348,0.000030992684,0.0015114206,0.000062720734,0.00018766498],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.988193,0.0038966548,0.0010410789,0.0031225563,0.0032321925,0.00051441387],"domain_scores_gemma":[0.47561035,0.474782,0.027756881,0.008548438,0.011557708,0.001744584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011660438,0.00073875167,0.0004556472,0.005520957,0.00048561848,0.002017575,0.0006430007,0.0011417934,0.0023439508],"category_scores_gemma":[0.14407578,0.00036785807,0.0010346299,0.0052366056,0.0012112994,0.0030123005,0.0012889814,0.0021960135,0.0007272408],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021675088,0.00011584779,0.95144165,0.00038523937,0.0003652434,0.00016168349,0.00070673175,0.0043726694,0.0012733983,0.00050656433,0.0008823264,0.03957183],"study_design_scores_gemma":[0.000009504388,0.00024425983,0.97290426,0.00012910586,0.00024932888,0.00031425682,0.00044173154,0.021850828,0.0011509209,0.0013452357,0.001324763,0.00003584504],"about_ca_topic_score_codex":0.003041037,"about_ca_topic_score_gemma":0.0031155725,"teacher_disagreement_score":0.011660438,"about_ca_system_score_codex":0.0007484779,"about_ca_system_score_gemma":0.00072288496,"threshold_uncertainty_score":0.061667025},"labels":[],"label_agreement":null},{"id":"W2963079908","doi":"10.1007/s10664-019-09743-4","title":"CAPS: a supervised technique for classifying Stack Overflow posts concerning API issues","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Queen's University","funders":"","keywords":"Computer science; Documentation; Application programming interface; Task (project management); Conditional random field; Readability; Field (mathematics); Baseline (sea); Data science; World Wide Web; Artificial intelligence; Machine learning; Engineering; Programming language","score_opus":0.0362815258937799,"score_gpt":0.3089535097084091,"score_spread":0.2726719838146292,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963079908","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37205988,0.0023793096,0.51544863,0.00083468785,0.0009740865,0.0024741306,0.02841619,0.0635778,0.013835296],"genre_scores_gemma":[0.4863953,0.0005604911,0.44510603,0.00041949027,0.0008418864,0.0013729015,0.046711177,0.0014534275,0.017139254],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9970414,0.0003988564,0.00037483155,0.00076101994,0.0011434715,0.0002802962],"domain_scores_gemma":[0.98970175,0.0037004002,0.0016152938,0.0013434131,0.003022974,0.000616258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022646796,0.0018874274,0.0011332909,0.010824432,0.0016282775,0.0015019494,0.0022615993,0.0020687603,0.0044256058],"category_scores_gemma":[0.0071877185,0.00044744337,0.0013070058,0.004916966,0.0007287445,0.002806585,0.001840433,0.0020537945,0.004101608],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009274927,0.0015421641,0.034193315,0.0009683235,0.00036460266,0.0006143465,0.00077339576,0.0049295845,0.04866533,0.0024416877,0.079598986,0.82498074],"study_design_scores_gemma":[0.00026146672,0.0011871892,0.06604637,0.00030956068,0.0006210195,0.0017550472,0.0018770604,0.7732438,0.087597266,0.010152664,0.056666207,0.00028233897],"about_ca_topic_score_codex":0.00590141,"about_ca_topic_score_gemma":0.01536159,"teacher_disagreement_score":0.010824432,"about_ca_system_score_codex":0.00066802814,"about_ca_system_score_gemma":0.0031385904,"threshold_uncertainty_score":0.014805138},"labels":[],"label_agreement":null},{"id":"W2963080927","doi":"10.4230/lipics.ecoop.2016.7","title":"Interprocedural Type Specialization of JavaScript Programs Without Type Analysis","year":2016,"lang":"en","type":"article","venue":"DROPS (Schloss Dagstuhl – Leibniz Center for Informatics)","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Programming language; JavaScript; Compiler; Software versioning; Block (permutation group theory); Simple (philosophy); Source code; Context (archaeology); Software","score_opus":0.020671696320407565,"score_gpt":0.2782466223932611,"score_spread":0.25757492607285354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963080927","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13413912,0.0005237037,0.8110126,0.00022635394,0.000083514715,0.00020794573,0.00042031356,0.046446573,0.006939932],"genre_scores_gemma":[0.6140419,0.0003122343,0.36785015,0.0003341751,0.00006762885,0.00013751161,0.001026574,0.009110563,0.007119247],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99621826,0.0010253083,0.0003367383,0.0007433583,0.0012646037,0.00041174554],"domain_scores_gemma":[0.9813442,0.0044121523,0.0013520791,0.011442441,0.0012501176,0.00019895342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025306433,0.0008105198,0.000893332,0.0011625283,0.0007672807,0.0017901073,0.0020960036,0.00059322733,0.002611543],"category_scores_gemma":[0.013815296,0.0008828099,0.0012486035,0.0015583409,0.001425239,0.003721005,0.0023944073,0.0017795848,0.0012592113],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001012083,0.00029857917,0.03500131,0.00078618224,0.00023056682,0.00041615256,0.0016745556,0.031409755,0.114552945,0.04162876,0.0068298965,0.7661592],"study_design_scores_gemma":[0.00016310319,0.0006645132,0.015795961,0.0002523548,0.00035651692,0.001451953,0.0002741249,0.560777,0.30944932,0.058434743,0.052172624,0.00020784485],"about_ca_topic_score_codex":0.002819659,"about_ca_topic_score_gemma":0.004226581,"teacher_disagreement_score":0.002819659,"about_ca_system_score_codex":0.0009993816,"about_ca_system_score_gemma":0.0021765293,"threshold_uncertainty_score":0.013383508},"labels":[],"label_agreement":null},{"id":"W2963223757","doi":"10.1109/icse-seip.2019.00026","title":"A Longitudinal Study of Identifying and Paying Down Architecture Debt","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Polytechnique Montréal","funders":"","keywords":"Technical debt; Maintainability; Debt; Architecture; Code refactoring; Computer science; Software architecture; Product (mathematics); Software; Business; Software engineering; Software development; Finance; Operating system","score_opus":0.0280041948498794,"score_gpt":0.29170501874710264,"score_spread":0.26370082389722327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963223757","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9983877,0.0000842122,0.0004551214,0.00033952927,0.000006712116,0.000038764607,0.00006067978,0.0000040110617,0.00062329485],"genre_scores_gemma":[0.9980532,0.00010595391,0.0005887266,0.0002303222,0.000009409802,0.00007764221,0.00009789626,0.0000056713825,0.0008311582],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99405766,0.0036637334,0.00033372972,0.00045997617,0.000871017,0.00061390596],"domain_scores_gemma":[0.9606065,0.01530586,0.009995616,0.003697341,0.0065245368,0.0038701233],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015528887,0.0003177374,0.00042007456,0.0017389645,0.0039368523,0.002517818,0.00072090223,0.0011321985,0.0025255915],"category_scores_gemma":[0.04144412,0.00074604317,0.00031824582,0.0012113191,0.0020027375,0.0034783746,0.0023089952,0.0029554344,0.0005713911],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020178917,0.0026355756,0.78459424,0.000060938157,0.000051845742,0.00072849664,0.1827697,0.00014411501,0.0010110071,0.0018864085,0.0013487701,0.024567066],"study_design_scores_gemma":[0.000053217864,0.003941235,0.72473556,0.00025413,0.00007164245,0.000836851,0.2541419,0.0018557013,0.001257528,0.002532649,0.010188507,0.00013110328],"about_ca_topic_score_codex":0.01050962,"about_ca_topic_score_gemma":0.01670927,"teacher_disagreement_score":0.015528887,"about_ca_system_score_codex":0.0018851514,"about_ca_system_score_gemma":0.0027924469,"threshold_uncertainty_score":0.082125604},"labels":[],"label_agreement":null},{"id":"W2963443047","doi":"10.5555/3021955.3021997","title":"Does Technical Debt Lead to the Rejection of Pull Requests","year":2016,"lang":"en","type":"article","venue":"IEEE International Conference on Cloud Computing Technology and Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Semtech (Canada)","funders":"","keywords":"Technical debt; Convention; Debt; Documentation; Identification (biology); Technical documentation; Computer science; Focus (optics); Risk analysis (engineering); Business; Accounting; Finance; Software; Law; Software development; Political science","score_opus":0.03509661961974468,"score_gpt":0.32513862062996346,"score_spread":0.29004200101021876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963443047","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9802886,0.0005253773,0.0037637493,0.0014641156,0.00007264405,0.00009337614,0.00020840131,0.00025293246,0.013330727],"genre_scores_gemma":[0.9958134,0.00017827668,0.0010545036,0.0003476139,0.00006747397,0.000042780564,0.00022745966,0.00009343797,0.0021750312],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98093975,0.006525272,0.0025672298,0.0016016782,0.005977112,0.002388931],"domain_scores_gemma":[0.6723823,0.17657506,0.1025204,0.014019958,0.02442632,0.010075965],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01496675,0.0005025635,0.00045640228,0.002420672,0.0016874158,0.0032928481,0.001228611,0.0020690055,0.008381016],"category_scores_gemma":[0.15385707,0.0004995817,0.00068860606,0.0016095341,0.0010752758,0.0047431756,0.002238639,0.0024545337,0.002168969],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077031803,0.00044499163,0.8940983,0.00055539986,0.000116926414,0.0015013549,0.014959096,0.0006794471,0.0060914424,0.00251691,0.0041906894,0.07407523],"study_design_scores_gemma":[0.0000744489,0.00070719945,0.89548194,0.00070738426,0.00022432492,0.0033544302,0.04779507,0.009253042,0.006204558,0.006230737,0.029805118,0.00016186388],"about_ca_topic_score_codex":0.0027059754,"about_ca_topic_score_gemma":0.0020915421,"teacher_disagreement_score":0.01496675,"about_ca_system_score_codex":0.0013417726,"about_ca_system_score_gemma":0.00195714,"threshold_uncertainty_score":0.0791527},"labels":[],"label_agreement":null},{"id":"W2963521149","doi":"10.1145/3304221.3325539","title":"Unexpected Tokens","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital","funders":"","keywords":"Computer science; Interpreter; Set (abstract data type); Process (computing); Syntax; Point (geometry); Programming language; Artificial intelligence","score_opus":0.0094676423964079,"score_gpt":0.2387978113532881,"score_spread":0.2293301689568802,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963521149","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026778819,0.0015121774,0.1843946,0.008913994,0.008471597,0.0014410294,0.02071857,0.023468105,0.7243012],"genre_scores_gemma":[0.26694322,0.0015210657,0.075746894,0.0051443647,0.0011170913,0.0011121214,0.019151196,0.011896631,0.6173674],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968437,0.00061464746,0.00034674635,0.00071252167,0.0009102138,0.00057216804],"domain_scores_gemma":[0.99357986,0.0018609893,0.00054268335,0.001865257,0.0016711066,0.00048002682],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0020834357,0.0012168904,0.0009902178,0.0020640837,0.0027656844,0.005513325,0.0033300663,0.0018817335,0.24233398],"category_scores_gemma":[0.01234855,0.00083093933,0.0011197482,0.0021087313,0.0018883869,0.010490922,0.0067314934,0.0021546143,0.10395642],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016728612,0.00013947373,0.0049718493,0.0014294366,0.00007570897,0.0027564203,0.0027935547,0.0013015938,0.005100864,0.40547606,0.3526258,0.22165637],"study_design_scores_gemma":[0.000053285912,0.00008935923,0.0010692223,0.00024315665,0.000037429305,0.001122378,0.0010893816,0.0018017275,0.0044765794,0.06054065,0.92940545,0.0000713888],"about_ca_topic_score_codex":0.0023860852,"about_ca_topic_score_gemma":0.0029022,"teacher_disagreement_score":0.757666,"about_ca_system_score_codex":0.002880692,"about_ca_system_score_gemma":0.0027188242,"threshold_uncertainty_score":0.8106879},"labels":[],"label_agreement":null},{"id":"W2963678029","doi":"10.1109/aire.2018.00007","title":"ELICA: An Automated Tool for Dynamic Extraction of Requirements Relevant Information","year":2018,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Requirements elicitation; Leverage (statistics); Conversation; Process (computing); Domain (mathematical analysis); Set (abstract data type); Information extraction; Flexibility (engineering); Software engineering; Requirements engineering; Knowledge management; Data science; Software; Artificial intelligence","score_opus":0.022066196148570354,"score_gpt":0.34886208100717275,"score_spread":0.3267958848586024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963678029","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0043621855,0.00010407112,0.8624765,0.00028109038,0.00006226157,0.00058138894,0.0050742957,0.12379169,0.0032663997],"genre_scores_gemma":[0.027027247,0.0001689136,0.9535324,0.00023903835,0.00003529587,0.0011172293,0.009559775,0.0047034444,0.0036166848],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99669766,0.0011281064,0.0003284057,0.0005565588,0.0011626966,0.00012660456],"domain_scores_gemma":[0.98381734,0.011452389,0.0011571661,0.0018242446,0.0015141242,0.00023474058],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047795707,0.0027330418,0.0010011876,0.0053489,0.0009936629,0.0021763833,0.00181436,0.0016324508,0.017982438],"category_scores_gemma":[0.01764982,0.0013214704,0.001584879,0.0017214458,0.0007870387,0.0027673682,0.0035275496,0.0021348586,0.010281826],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006522785,0.00043775683,0.0036922367,0.0034760428,0.0002378951,0.0019677363,0.004084557,0.016871579,0.09848344,0.022562888,0.15235174,0.6951818],"study_design_scores_gemma":[0.00031142606,0.00036239132,0.004077052,0.00065876986,0.00016425858,0.0024605996,0.0020524485,0.4727162,0.109766446,0.045530077,0.36151314,0.0003872634],"about_ca_topic_score_codex":0.0017512802,"about_ca_topic_score_gemma":0.0032648535,"teacher_disagreement_score":0.017982438,"about_ca_system_score_codex":0.00089023943,"about_ca_system_score_gemma":0.002692051,"threshold_uncertainty_score":0.0601573},"labels":[],"label_agreement":null},{"id":"W2963739263","doi":"","title":"Neural Guided Constraint Logic Programming for Program Synthesis","year":2018,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Artificial neural network; Artificial intelligence; Constraint programming; Concurrent constraint logic programming; Constraint (computer-aided design); Inductive programming; Logic programming; Constraint logic programming; Recurrent neural network; Constraint satisfaction; Representation (politics); Theoretical computer science; Programming paradigm; Programming language; Mathematical optimization; Mathematics; Stochastic programming","score_opus":0.11624863787366077,"score_gpt":0.24186302803171036,"score_spread":0.1256143901580496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963739263","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00558528,0.0002301891,0.9886348,0.00023730726,0.000020807387,0.000045633486,0.000072135656,0.00073192257,0.0044419672],"genre_scores_gemma":[0.17978995,0.00042004188,0.8143584,0.00022794047,0.000030175455,0.0002825398,0.0002534118,0.00035506263,0.0042824764],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9994709,0.00017516266,0.000030321218,0.00011206379,0.00016649006,0.00004520115],"domain_scores_gemma":[0.9990687,0.0006710522,0.000056404588,0.000089602305,0.000091144495,0.000023138951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007657683,0.0007902334,0.00047979754,0.00047129925,0.00036941897,0.00075219857,0.001105595,0.0008057768,0.0052433186],"category_scores_gemma":[0.003307683,0.00044468633,0.00081758486,0.00062312814,0.0010367285,0.0012946423,0.00089335244,0.0018520296,0.00071278657],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000067080175,0.00006599297,0.00024740442,0.00028917525,0.000045953886,0.00008575022,0.000090409674,0.7249474,0.005471263,0.13093674,0.003179622,0.13457316],"study_design_scores_gemma":[0.000012191464,0.000015999767,0.000020431387,0.000019686468,0.000007897974,0.000018569242,0.0000083626,0.9379011,0.0021764154,0.056626596,0.003187057,0.0000056925173],"about_ca_topic_score_codex":0.002896416,"about_ca_topic_score_gemma":0.0071538454,"teacher_disagreement_score":0.0052433186,"about_ca_system_score_codex":0.0010952658,"about_ca_system_score_gemma":0.001457864,"threshold_uncertainty_score":0.017540693},"labels":[],"label_agreement":null},{"id":"W2964084063","doi":"10.4230/lipics.ecoop.2015.101","title":"Simple and Effective Type Check Removal through Lazy Basic Block Versioning","year":2015,"lang":"en","type":"article","venue":"DROPS (Schloss Dagstuhl – Leibniz Center for Informatics)","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Programming language; Software versioning; Compiler; Just-in-time compilation; JavaScript; Type inference; Python (programming language); Toolchain; Parallel computing; Software; Inference","score_opus":0.024121287491206147,"score_gpt":0.2821698156898441,"score_spread":0.258048528198638,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964084063","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09032282,0.0005337058,0.8584301,0.00015956056,0.00011758786,0.00016183603,0.0002795529,0.04620769,0.0037871534],"genre_scores_gemma":[0.47820312,0.0004258103,0.5034189,0.00024994442,0.00007703648,0.00018642598,0.00080804597,0.011905557,0.0047252015],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984157,0.00028840557,0.00016095249,0.0002541309,0.00067103904,0.00020975411],"domain_scores_gemma":[0.99490225,0.0013029223,0.000508986,0.002602808,0.00056008704,0.00012295926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012452743,0.0006664536,0.0006227158,0.00078590447,0.00051167374,0.0013188708,0.0015753936,0.00048663866,0.0027845637],"category_scores_gemma":[0.0060997647,0.00072950096,0.0008733806,0.0007249301,0.0009936871,0.0023977456,0.0016249446,0.00143763,0.0015494373],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096087507,0.0002924692,0.020551814,0.0008150927,0.0001733808,0.0006508691,0.0010295338,0.041444857,0.2870661,0.033512037,0.011644361,0.6018586],"study_design_scores_gemma":[0.00022767094,0.0005042346,0.0086745,0.00023310467,0.0002475985,0.0014222631,0.00015463465,0.39736715,0.49248904,0.03868454,0.059763663,0.00023163145],"about_ca_topic_score_codex":0.0014041148,"about_ca_topic_score_gemma":0.0017422339,"teacher_disagreement_score":0.0027845637,"about_ca_system_score_codex":0.00045016533,"about_ca_system_score_gemma":0.001873678,"threshold_uncertainty_score":0.009315252},"labels":[],"label_agreement":null},{"id":"W2964175311","doi":"10.1109/tse.2018.2868349","title":"Debugging Static Analysis","year":2018,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Universität Paderborn; École Polytechnique Fédérale de Lausanne; Heinz Nixdorf Stiftung; Deutsche Forschungsgemeinschaft","keywords":"Static analysis; Debugging; Computer science; Debugger; Static program analysis; Programming language; Program analysis; Algorithmic program debugging; Software bug; Source code; Data-flow analysis; Call graph; Code (set theory); Tracing; Taint checking; Software engineering; Data flow diagram; Database; Software; Software development","score_opus":0.013345475541958063,"score_gpt":0.25049080371884513,"score_spread":0.23714532817688708,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964175311","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07586462,0.0021104568,0.8193345,0.002171436,0.00072713714,0.0005529908,0.0032068528,0.061371952,0.03466006],"genre_scores_gemma":[0.5319442,0.0017231868,0.4349337,0.0010973373,0.00026332805,0.00042130981,0.005686294,0.009968057,0.013962522],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9883937,0.0042959102,0.00093524257,0.0019887574,0.0037346806,0.00065156695],"domain_scores_gemma":[0.9344354,0.031865686,0.00394537,0.01234233,0.01662031,0.0007909361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012373361,0.002106892,0.0010445262,0.005723947,0.0015190635,0.0030419189,0.0018104868,0.0012164471,0.011575365],"category_scores_gemma":[0.06857033,0.0010182904,0.0008696718,0.0029602393,0.0009999969,0.0057580895,0.0027174861,0.0017754643,0.0059984955],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077093655,0.00029055093,0.03547023,0.0012904486,0.00015620772,0.00072383846,0.008347537,0.008312818,0.02270024,0.028518658,0.065960094,0.82745844],"study_design_scores_gemma":[0.00021848634,0.0007992882,0.027574575,0.0023363638,0.0005216845,0.0033632726,0.0061707483,0.1314228,0.0955759,0.07363362,0.6579229,0.00046036957],"about_ca_topic_score_codex":0.0027682881,"about_ca_topic_score_gemma":0.0030214149,"teacher_disagreement_score":0.012373361,"about_ca_system_score_codex":0.0010700546,"about_ca_system_score_gemma":0.003261547,"threshold_uncertainty_score":0.06543738},"labels":[],"label_agreement":null},{"id":"W2964177436","doi":"10.1145/3345629.3345632","title":"On Usefulness of the Deep-Learning-Based Bug Localization Models to Practitioners","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software bug; Software; Predictive modelling; Debugging; Security bug","score_opus":0.020609452413016257,"score_gpt":0.24869465598200433,"score_spread":0.22808520356898807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964177436","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38412097,0.010970663,0.51943934,0.04768191,0.00065763324,0.00017067006,0.001023583,0.0030743536,0.032860864],"genre_scores_gemma":[0.95339185,0.0014858246,0.041412007,0.0009881814,0.00015922968,0.000073091134,0.0004501086,0.000117878364,0.0019219307],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99661213,0.0020607864,0.00011054268,0.00048012464,0.0005908141,0.00014558039],"domain_scores_gemma":[0.94567084,0.041858114,0.001863401,0.0028092111,0.006783351,0.0010151928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009034275,0.0012967446,0.0005764895,0.001456957,0.000534213,0.0019456374,0.0010492866,0.0020530322,0.0035129753],"category_scores_gemma":[0.07034483,0.00039367768,0.0003860921,0.0010159401,0.0011006476,0.004035171,0.001510338,0.00258059,0.0007198375],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094037486,0.0007800169,0.05536065,0.0009125932,0.00020758584,0.00033769445,0.0013895797,0.31051838,0.0027801206,0.054240804,0.041618586,0.53091353],"study_design_scores_gemma":[0.00004890616,0.00016622298,0.0023669237,0.00021724268,0.000047456295,0.000074485775,0.00021217868,0.9484244,0.0010845407,0.04389028,0.003445089,0.000022344368],"about_ca_topic_score_codex":0.009042343,"about_ca_topic_score_gemma":0.0064847507,"teacher_disagreement_score":0.009042343,"about_ca_system_score_codex":0.0018116315,"about_ca_system_score_gemma":0.0022937385,"threshold_uncertainty_score":0.047778368},"labels":[],"label_agreement":null},{"id":"W2964301819","doi":"","title":"Building an OSS Quality Estimation Model with CATREG","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Categorical variable; Software; Computer science; Regression analysis; Quality (philosophy); Software quality; Data mining; Estimation; Process (computing); Linear regression; Regression; Statistics; Software development; Machine learning; Mathematics; Engineering; Systems engineering","score_opus":0.09189924603784294,"score_gpt":0.3719724061074943,"score_spread":0.2800731600696514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964301819","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040559955,0.00012568738,0.956266,0.00013347135,0.000015615495,0.000068523834,0.0002993769,0.0015024523,0.0010289968],"genre_scores_gemma":[0.69191545,0.0003044686,0.3012525,0.0000957872,0.000029447092,0.0006134103,0.0016767964,0.0003051558,0.003806899],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99832875,0.00061466446,0.000111603804,0.0004403946,0.00038421276,0.00012036767],"domain_scores_gemma":[0.99546576,0.0025806886,0.0006186869,0.00027872922,0.0009963238,0.000059669263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030387985,0.0009520262,0.000891491,0.002005281,0.00031881765,0.0014440769,0.0015209805,0.0010130738,0.0027808219],"category_scores_gemma":[0.008447999,0.0006202225,0.001585359,0.0016822523,0.00047753533,0.0012770735,0.0010190367,0.0012072306,0.0011271691],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001035364,0.00011548106,0.012814794,0.00012926618,0.00013821563,0.00018995843,0.0001844212,0.90496856,0.0026750793,0.011749376,0.0011128611,0.06581846],"study_design_scores_gemma":[0.000003130012,0.000020690197,0.0006861098,0.0000067998426,0.000014866464,0.000022050155,0.000012317845,0.99666303,0.00032886755,0.0018468649,0.0003889216,0.000006412465],"about_ca_topic_score_codex":0.008984316,"about_ca_topic_score_gemma":0.0061990474,"teacher_disagreement_score":0.008984316,"about_ca_system_score_codex":0.001036809,"about_ca_system_score_gemma":0.0012020395,"threshold_uncertainty_score":0.017864048},"labels":[],"label_agreement":null},{"id":"W2965888979","doi":"10.22215/etd/2018-12853","title":"Leveraging Insights from Mobile App Reviews to Support Release Planning and Maintenance","year":2018,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Mobile apps; Key (lock); Process (computing); World Wide Web; Space (punctuation); App store; Data science; Plan (archaeology); Population; Computer security","score_opus":0.025037627961720408,"score_gpt":0.3031684398876298,"score_spread":0.2781308119259094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2965888979","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58989745,0.014864556,0.33297586,0.006155633,0.0013564493,0.0015478971,0.019444548,0.0048379097,0.028919714],"genre_scores_gemma":[0.7610167,0.005525168,0.20863327,0.00054626155,0.0012039642,0.0005377718,0.011837666,0.00068531185,0.010013893],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99668545,0.0010577912,0.00026687412,0.0005075221,0.0013513248,0.00013107016],"domain_scores_gemma":[0.9674172,0.01763295,0.004218653,0.001108291,0.00912115,0.0005017082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030779336,0.000994965,0.0006122531,0.005956592,0.000550631,0.0023452255,0.00043076277,0.0005418205,0.0014863068],"category_scores_gemma":[0.030366426,0.00041815542,0.00047824712,0.0033249331,0.00021306687,0.003089093,0.0009017319,0.0009481573,0.0016748513],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053567707,0.00030526126,0.07338092,0.002440473,0.00035518376,0.0009317674,0.0081672445,0.0043379273,0.022522362,0.0023983906,0.042522967,0.8421019],"study_design_scores_gemma":[0.00011078015,0.0012190205,0.4122211,0.0018911484,0.0010362429,0.0022046808,0.009729382,0.2916876,0.03486947,0.018558858,0.22602656,0.0004451128],"about_ca_topic_score_codex":0.002874177,"about_ca_topic_score_gemma":0.011162257,"teacher_disagreement_score":0.005956592,"about_ca_system_score_codex":0.0006110901,"about_ca_system_score_gemma":0.00089601206,"threshold_uncertainty_score":0.01627791},"labels":[],"label_agreement":null},{"id":"W2966745404","doi":"10.1109/mobilesoft.2019.00021","title":"A Comparison of Bugs Across the iOS and Android Platforms of Two Open Source Cross Platform Browser Apps","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Waterloo","funders":"","keywords":"Android (operating system); Computer science; World Wide Web; Mobile apps; Open source; Mobile device; Android app; Operating system; Software","score_opus":0.039530811225194064,"score_gpt":0.3747698814468279,"score_spread":0.3352390702216338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966745404","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99279964,0.00096729444,0.002426444,0.00025740208,0.000052847594,0.0001194857,0.0007714625,0.000288084,0.0023172079],"genre_scores_gemma":[0.9922077,0.0004020012,0.003911596,0.000082479586,0.00004958377,0.00012731379,0.0018449866,0.00015275716,0.001221579],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9888293,0.0027814973,0.0020682411,0.0016954569,0.0038836014,0.00074189453],"domain_scores_gemma":[0.7091231,0.20569032,0.050700184,0.0071943398,0.02388161,0.003410477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008641065,0.00059656496,0.00052433665,0.0080030095,0.00078302564,0.0023568447,0.0007557397,0.0011971627,0.0013749034],"category_scores_gemma":[0.08947732,0.00050755724,0.0011966496,0.0036369283,0.0011072932,0.004934898,0.0016578804,0.001279756,0.00040878908],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012440981,0.0005594466,0.8325867,0.0014945403,0.00038401698,0.00105738,0.022283059,0.0014607215,0.0067843627,0.0013121511,0.0032509086,0.12758264],"study_design_scores_gemma":[0.000026798554,0.0011332071,0.97379535,0.00029208104,0.00024433766,0.001294228,0.0105539225,0.005371923,0.0020668444,0.0006748239,0.0044623557,0.000084210515],"about_ca_topic_score_codex":0.0032567559,"about_ca_topic_score_gemma":0.004832984,"teacher_disagreement_score":0.008641065,"about_ca_system_score_codex":0.00079423335,"about_ca_system_score_gemma":0.0007148185,"threshold_uncertainty_score":0.04569888},"labels":[],"label_agreement":null},{"id":"W2966845482","doi":"10.1109/mise.2019.00009","title":"Detecting Emergent Behaviors and Implied Scenarios in Scenario-Based Specifications: A Machine Learning Approach","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Software deployment; Distributed computing; Scale (ratio); Software system; Software; Reliability engineering; Software engineering; Engineering","score_opus":0.03056539758250197,"score_gpt":0.2558895466680578,"score_spread":0.2253241490855558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966845482","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.077403925,0.00014806206,0.91842353,0.00037595214,0.000017221679,0.00021843602,0.00050202996,0.0019790353,0.0009318563],"genre_scores_gemma":[0.63206875,0.00010394538,0.364891,0.00016379796,0.000027443073,0.00023715226,0.0018032647,0.00006158307,0.0006430514],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983699,0.00059241743,0.00017442938,0.00040091507,0.00033547293,0.00012686342],"domain_scores_gemma":[0.99012494,0.0069241035,0.0010447371,0.0006979734,0.0009788158,0.00022941947],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016880825,0.0011902875,0.00056729245,0.002095822,0.00045793556,0.0010815009,0.0016394649,0.0013018735,0.0011335474],"category_scores_gemma":[0.0075830305,0.00048296203,0.0011682949,0.0010219403,0.00062142924,0.0018217813,0.00086817663,0.0015273686,0.00033091783],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018837096,0.0005912995,0.016974524,0.00017009884,0.00015237907,0.00042076223,0.00024607504,0.7938248,0.0059928917,0.0059310533,0.0016308221,0.17387696],"study_design_scores_gemma":[0.0000040874693,0.000019051029,0.00043315097,0.000005245502,0.000006222022,0.000028524211,0.000026003734,0.99573356,0.00062421494,0.0029422303,0.00017178086,0.0000058777887],"about_ca_topic_score_codex":0.005438478,"about_ca_topic_score_gemma":0.006795571,"teacher_disagreement_score":0.005438478,"about_ca_system_score_codex":0.0010473062,"about_ca_system_score_gemma":0.0012625307,"threshold_uncertainty_score":0.0108136535},"labels":[],"label_agreement":null},{"id":"W2967302497","doi":"10.18293/seke2019-219","title":"Feature Evaluation for Automatic Bug Report Summarization (S)","year":2019,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Precision and recall; Feature (linguistics); Natural language processing; Recall; Artificial intelligence; Sentence; Software; Software bug; Software regression; Information retrieval; Software development; Programming language; Software quality; Linguistics","score_opus":0.02174730908305818,"score_gpt":0.271068682737073,"score_spread":0.24932137365401483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967302497","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46418405,0.005520133,0.41909537,0.0014114408,0.0005955126,0.001492499,0.022646967,0.0801787,0.004875344],"genre_scores_gemma":[0.714016,0.00070086436,0.24462985,0.00016352748,0.0002465885,0.00090522255,0.035800498,0.0009777276,0.0025597329],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99372935,0.0018780215,0.00093733723,0.0011988817,0.0018702467,0.00038605102],"domain_scores_gemma":[0.95815665,0.022611633,0.0045847977,0.002972726,0.011077487,0.0005966585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008816846,0.00210495,0.0016579956,0.009027217,0.00067973376,0.0021554108,0.0017452319,0.0011703856,0.002096594],"category_scores_gemma":[0.049396098,0.0004718362,0.001625224,0.004205985,0.00029251567,0.0025144517,0.0011736791,0.0011729923,0.0017971882],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017875102,0.00074791774,0.07169203,0.0025941175,0.00075633,0.00065151363,0.0011312236,0.01909544,0.041772515,0.0009820606,0.0413879,0.81740147],"study_design_scores_gemma":[0.00043400482,0.0027170428,0.13148479,0.00048789912,0.0017017436,0.0013998181,0.0012616314,0.7375469,0.08948749,0.002716908,0.030395003,0.00036673713],"about_ca_topic_score_codex":0.0061535253,"about_ca_topic_score_gemma":0.005581913,"teacher_disagreement_score":0.009027217,"about_ca_system_score_codex":0.0010434138,"about_ca_system_score_gemma":0.0013027777,"threshold_uncertainty_score":0.046628535},"labels":[],"label_agreement":null},{"id":"W2967307245","doi":"10.1002/smr.2208","title":"A delta‐oriented approach to support the safe reuse of black‐box code rewriters","year":2019,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Centre National de la Recherche Scientifique","keywords":"Computer science; Rewriting; Reuse; Programming language; Forward chaining; Chaining; Android (operating system); Black box; Code (set theory); Code reuse; Software engineering; Software; Operating system; Artificial intelligence; Set (abstract data type)","score_opus":0.01514872626915628,"score_gpt":0.26820557456062716,"score_spread":0.2530568482914709,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967307245","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010382099,0.00008158688,0.98028255,0.00016564406,0.000056229994,0.00013791856,0.00006934226,0.0076392177,0.0011855417],"genre_scores_gemma":[0.18619116,0.00016165338,0.8056402,0.0003185402,0.000051917857,0.0003119779,0.0003641149,0.0022315222,0.0047288747],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99579406,0.0010941392,0.00061832636,0.0007305313,0.0013701388,0.00039282755],"domain_scores_gemma":[0.9824265,0.004323577,0.0011767431,0.009286739,0.0022235946,0.00056276046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006141543,0.0009411034,0.0006921303,0.0020653845,0.00085306197,0.0026208067,0.004596119,0.0019929942,0.0032261682],"category_scores_gemma":[0.015679162,0.0015226231,0.00146818,0.0007858943,0.0029730264,0.0045095338,0.0048404434,0.0031728938,0.0013044253],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009235101,0.0009921129,0.016931713,0.0010925228,0.00028037632,0.0024773995,0.0040019616,0.11990086,0.07685164,0.35391477,0.010943246,0.41168994],"study_design_scores_gemma":[0.00014202323,0.00035436664,0.00092735124,0.00028896742,0.00018028567,0.0010758104,0.0002734162,0.6868982,0.083874345,0.16472982,0.061104525,0.00015092925],"about_ca_topic_score_codex":0.0028279584,"about_ca_topic_score_gemma":0.0033791375,"teacher_disagreement_score":0.006141543,"about_ca_system_score_codex":0.001061633,"about_ca_system_score_gemma":0.0023342674,"threshold_uncertainty_score":0.03248},"labels":[],"label_agreement":null},{"id":"W2967325635","doi":"10.1145/3338906.3338944","title":"Bisecting commits and modeling commit risk during testing","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Commit; Computer science; Reliability engineering; Acceptance testing; Test strategy; Test Management Approach; Test case; Non-regression testing; Risk-based testing; Software; Random testing; Software system; Operating system; Database; Engineering; Software engineering; Machine learning; Software construction","score_opus":0.03198539147721205,"score_gpt":0.2542237448922663,"score_spread":0.22223835341505424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967325635","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5166421,0.0013190468,0.46836302,0.0014337767,0.00012108786,0.00071247894,0.0009616952,0.0013605688,0.009086172],"genre_scores_gemma":[0.9412245,0.00032935556,0.053842973,0.00009061131,0.000044975197,0.00034915438,0.00054293574,0.00017892157,0.0033965784],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912765,0.003071988,0.0005842095,0.001690259,0.0024240038,0.0009530089],"domain_scores_gemma":[0.9090166,0.06405255,0.01448499,0.0057822797,0.004641009,0.0020225116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011458141,0.0016291251,0.0009591111,0.0034110814,0.00083791354,0.0030732844,0.0034965286,0.0029594514,0.003595218],"category_scores_gemma":[0.08817189,0.0015874979,0.0013550016,0.0020181166,0.002487031,0.004695746,0.0022134979,0.00291943,0.0005069712],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001693853,0.00012137639,0.017021151,0.000076578806,0.000059332495,0.00030475023,0.00056416664,0.94064116,0.000976323,0.022119537,0.00065117877,0.017294962],"study_design_scores_gemma":[0.000036637302,0.00013803045,0.0026029178,0.000036177502,0.00004135645,0.00011504907,0.000099448334,0.97374207,0.00060751854,0.021687604,0.0008684638,0.000024616482],"about_ca_topic_score_codex":0.01298507,"about_ca_topic_score_gemma":0.008009732,"teacher_disagreement_score":0.01298507,"about_ca_system_score_codex":0.0030602803,"about_ca_system_score_gemma":0.002790652,"threshold_uncertainty_score":0.06059718},"labels":[],"label_agreement":null},{"id":"W2967758098","doi":"10.1145/3338906.3341182","title":"CloneCognition: machine learning based code clone validation tool","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"clone (Java method); Computer science; Source code; Code (set theory); Generalization; Software maintenance; Programming language; Program comprehension; Software; Codebase; Software system; Process (computing); Artificial intelligence; Software engineering; Machine learning; Set (abstract data type)","score_opus":0.017681970754895488,"score_gpt":0.2522891837512765,"score_spread":0.23460721299638101,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967758098","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026393568,0.0006541056,0.61464983,0.00033702792,0.00019093357,0.0003438168,0.0023282408,0.35244638,0.0026560668],"genre_scores_gemma":[0.23611481,0.00043699547,0.7303419,0.0007274958,0.00008219555,0.0007833675,0.012432072,0.011998499,0.007082683],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99787223,0.00032024368,0.00019005612,0.00063382194,0.00085723563,0.00012641055],"domain_scores_gemma":[0.99325424,0.0038269716,0.0007429561,0.0009298455,0.001093829,0.00015208972],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018985268,0.0015600036,0.00080593274,0.003337286,0.0005933897,0.0012732557,0.0029751004,0.0021028356,0.0046895524],"category_scores_gemma":[0.014610486,0.0007033706,0.0011818846,0.0012316569,0.00083053246,0.0024670297,0.0018049603,0.001922761,0.002986886],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005934171,0.0004140653,0.018502207,0.0010445814,0.00026171593,0.001155311,0.00067190855,0.028705085,0.040919285,0.0063598035,0.09777336,0.80359936],"study_design_scores_gemma":[0.00017147973,0.00029265817,0.0074750595,0.00023221616,0.000099842066,0.0015678813,0.00017061518,0.79996634,0.117642306,0.0117750745,0.060450878,0.00015562878],"about_ca_topic_score_codex":0.0027584617,"about_ca_topic_score_gemma":0.003031985,"teacher_disagreement_score":0.0046895524,"about_ca_system_score_codex":0.00076614745,"about_ca_system_score_gemma":0.0014610577,"threshold_uncertainty_score":0.015688062},"labels":[],"label_agreement":null},{"id":"W2967780468","doi":"10.1007/s10664-019-09736-3","title":"The impact of context metrics on just-in-time defect prediction","year":2019,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada","keywords":"Source lines of code; Context (archaeology); Metric (unit); Computer science; Measure (data warehouse); Machine learning; Software metric; Software; Data mining; Artificial intelligence; Statistics; Software quality; Mathematics; Software development; Programming language; Engineering; Operations management","score_opus":0.020471050761877507,"score_gpt":0.29267956045619137,"score_spread":0.27220850969431387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967780468","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9762323,0.0017151366,0.017995195,0.00050771236,0.0001448427,0.000033031472,0.0010686578,0.0005726783,0.0017303476],"genre_scores_gemma":[0.9952567,0.00009406863,0.003887393,0.00002570634,0.000034071483,0.000007530691,0.0005017814,0.000033190234,0.00015949813],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9933216,0.0031915302,0.0005431055,0.0012094405,0.0013324588,0.0004017906],"domain_scores_gemma":[0.83823,0.12933287,0.009538211,0.010793479,0.009162536,0.0029429155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073770573,0.0011225974,0.00084540615,0.002664894,0.00048980577,0.0017082756,0.00086500996,0.0011719442,0.0010298378],"category_scores_gemma":[0.0904149,0.00021368168,0.0005335115,0.0020630793,0.0005178119,0.004036544,0.0011877618,0.0014589666,0.0003470936],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013092734,0.0007091391,0.79001445,0.00020198133,0.00034526055,0.00018838403,0.00021994098,0.03706287,0.002314791,0.0012821052,0.0027647384,0.16358714],"study_design_scores_gemma":[0.000086557455,0.002515583,0.32797652,0.00013142591,0.00040948132,0.0006979722,0.00053259614,0.6522227,0.00506515,0.008336649,0.0019243264,0.00010103265],"about_ca_topic_score_codex":0.003641223,"about_ca_topic_score_gemma":0.0072049634,"teacher_disagreement_score":0.0073770573,"about_ca_system_score_codex":0.00049118756,"about_ca_system_score_gemma":0.00096017716,"threshold_uncertainty_score":0.0390141},"labels":[],"label_agreement":null},{"id":"W2967794585","doi":"10.1109/re.2019.00017","title":"A Machine Learning-Based Approach for Demarcating Requirements in Textual Specifications","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds National de la Recherche Luxembourg; European Commission","keywords":"Computer science; Artificial intelligence; Natural language processing; Programming language; Software engineering","score_opus":0.06556066741738399,"score_gpt":0.29487738119851675,"score_spread":0.22931671378113277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967794585","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06622251,0.0021237656,0.87047714,0.0012122667,0.00021312728,0.0009939666,0.011452046,0.038690828,0.008614349],"genre_scores_gemma":[0.17552629,0.0003946056,0.7793609,0.00058350194,0.000069046415,0.00066659274,0.03828906,0.00040296413,0.004707073],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950321,0.0012737245,0.00065161294,0.0015830216,0.0012430067,0.00021664004],"domain_scores_gemma":[0.9909594,0.004239467,0.0009654329,0.0011786728,0.0024707068,0.00018622227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030143673,0.0018717749,0.0007271698,0.006737324,0.0009056506,0.0018949069,0.0024600565,0.0019494216,0.002823684],"category_scores_gemma":[0.011190581,0.00046742673,0.0016747414,0.0033396515,0.0006612363,0.002325727,0.0014782139,0.0024473742,0.0037227129],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032779813,0.0005898799,0.011947782,0.0009578504,0.00016920663,0.0004907675,0.00054088904,0.034407426,0.017949538,0.0052243974,0.043573406,0.8838211],"study_design_scores_gemma":[0.00010979644,0.0002613091,0.008021209,0.00027154543,0.00012932261,0.0008710304,0.0005848488,0.89658654,0.030664429,0.013433203,0.048980277,0.00008640199],"about_ca_topic_score_codex":0.012114855,"about_ca_topic_score_gemma":0.022618795,"teacher_disagreement_score":0.012114855,"about_ca_system_score_codex":0.0019123508,"about_ca_system_score_gemma":0.0026104988,"threshold_uncertainty_score":0.02408868},"labels":[],"label_agreement":null},{"id":"W2968094109","doi":"10.1145/3450952","title":"Bidirectional Typing","year":2021,"lang":"en","type":"preprint","venue":"ACM Computing Surveys","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Typing; Type inference; Computer science; Undecidable problem; Inference; Programming language; Annotation; Type (biology); Theoretical computer science; Artificial intelligence; Biology; Speech recognition; Decidability","score_opus":0.04907027266529844,"score_gpt":0.308312607897336,"score_spread":0.25924233523203755,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2968094109","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007302192,0.0008445396,0.92930526,0.0013558776,0.000775858,0.00019500808,0.001240122,0.007058833,0.051922336],"genre_scores_gemma":[0.30595785,0.0023641812,0.5795293,0.003347197,0.0008044278,0.0009067173,0.0033580454,0.009431893,0.09430048],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98918706,0.0029523217,0.0012885119,0.0023031516,0.003118566,0.001150381],"domain_scores_gemma":[0.9773951,0.005790675,0.0009998525,0.010877814,0.0043517146,0.0005848176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008150815,0.0012431755,0.0013849279,0.0021170015,0.0028952167,0.006434172,0.0034326352,0.0023006857,0.023454402],"category_scores_gemma":[0.022997478,0.0017349924,0.0021447712,0.0026151303,0.0044691456,0.012999331,0.009434218,0.004373391,0.012019297],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014033282,0.000042061765,0.002057688,0.00037473263,0.000057205576,0.00020740322,0.0013493327,0.0013559206,0.0033276768,0.8798891,0.019675095,0.0915235],"study_design_scores_gemma":[0.000043148448,0.000055331733,0.00041346432,0.00039811633,0.00013321138,0.00068814476,0.00043384594,0.008580934,0.011455731,0.58926123,0.38844746,0.000089361674],"about_ca_topic_score_codex":0.0028673497,"about_ca_topic_score_gemma":0.0035957596,"teacher_disagreement_score":0.023454402,"about_ca_system_score_codex":0.0017925074,"about_ca_system_score_gemma":0.0036887412,"threshold_uncertainty_score":0.07846278},"labels":[],"label_agreement":null},{"id":"W2970646592","doi":"10.1007/978-3-030-29238-6_10","title":"Empirical Analysis of Object-Oriented Metrics and Centrality Measures for Predicting Fault-Prone Classes in Object-Oriented Software","year":2019,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Centrality; Computer science; Software metric; Software; Data mining; Object-oriented programming; Object (grammar); Software fault tolerance; Java; Software development; Artificial intelligence; Software quality; Programming language; Mathematics","score_opus":0.05846315919585426,"score_gpt":0.3387307878319006,"score_spread":0.2802676286360463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970646592","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9809276,0.0008232716,0.015758017,0.00018607294,0.000018260147,0.000014204494,0.00025007376,0.00005796381,0.0019644664],"genre_scores_gemma":[0.9930011,0.00020779137,0.005689885,0.000012976238,0.0000304001,0.000014732583,0.00048527433,0.000023586377,0.0005342181],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99781287,0.001129444,0.00010595578,0.00023108245,0.0006106258,0.00010991065],"domain_scores_gemma":[0.84018004,0.14947124,0.0038604625,0.0024886315,0.0032163449,0.00078331487],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051737083,0.0004939121,0.00030639465,0.0035154435,0.0004051553,0.0011582026,0.00095799466,0.00064394856,0.00151034],"category_scores_gemma":[0.049969863,0.00019887438,0.0003716764,0.002869151,0.00076886057,0.0022635574,0.00067037094,0.0011179295,0.0003352798],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004004284,0.00061218324,0.82062,0.0001211861,0.00020820879,0.00013544929,0.0009291652,0.025741622,0.0016453788,0.007975475,0.002288082,0.13932292],"study_design_scores_gemma":[0.000027520804,0.00048015654,0.55421203,0.00008732141,0.00014825597,0.0005158746,0.001443961,0.42227685,0.0018940439,0.017547509,0.0013140483,0.00005236665],"about_ca_topic_score_codex":0.0025575869,"about_ca_topic_score_gemma":0.0038613142,"teacher_disagreement_score":0.0051737083,"about_ca_system_score_codex":0.00069774216,"about_ca_system_score_gemma":0.0004736335,"threshold_uncertainty_score":0.027361512},"labels":[],"label_agreement":null},{"id":"W2971805364","doi":"10.48550/arxiv.2103.03506","title":"Does chronology matter in JIT defect prediction? A Partial Replication Study","year":2021,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Eclipse; Brier score; Wilcoxon signed-rank test; Code (set theory); Replication (statistics); Data mining; Statistics; Artificial intelligence; Set (abstract data type); Mathematics; Programming language","score_opus":0.0369766811427843,"score_gpt":0.20261096639095677,"score_spread":0.16563428524817247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2971805364","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9553434,0.0082623055,0.023481112,0.001978136,0.0005488306,0.0004035029,0.005230047,0.00059759425,0.004155045],"genre_scores_gemma":[0.99093926,0.00050391274,0.0040636617,0.00041087894,0.00017097154,0.00017661558,0.0024959245,0.00026929297,0.00096946635],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97279817,0.0140226465,0.0023711424,0.007311584,0.0028248224,0.0006715322],"domain_scores_gemma":[0.56968874,0.2582159,0.03823818,0.091324374,0.039580785,0.002952031],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.068753116,0.0010421417,0.0012946116,0.0023406947,0.0012444154,0.0026139813,0.0025383208,0.0016849274,0.0043657986],"category_scores_gemma":[0.27329007,0.00081097597,0.003552592,0.0028681664,0.0014902835,0.005571732,0.0017850888,0.0021426487,0.0017798133],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018330949,0.0005218019,0.9197436,0.00096208224,0.0023356262,0.0005034069,0.004180006,0.0033135687,0.0017388802,0.001593244,0.005585529,0.05768919],"study_design_scores_gemma":[0.00046922764,0.0031488817,0.8958704,0.0011297439,0.003966637,0.0015919422,0.0047549484,0.050090414,0.004610551,0.007820562,0.02625497,0.00029170606],"about_ca_topic_score_codex":0.018321162,"about_ca_topic_score_gemma":0.010499446,"teacher_disagreement_score":0.9312469,"about_ca_system_score_codex":0.0011567516,"about_ca_system_score_gemma":0.0022647153,"threshold_uncertainty_score":0.36360556},"labels":[{"model":"gemma","categories":["metaresearch"],"domain":"reproducibility","study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":["metaresearch"],"domain":"reproducibility","study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"agree"},{"id":"W2972291021","doi":"10.1016/j.jss.2019.110407","title":"An empirical study on bug propagation through code cloning","year":2019,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"University of Saskatchewan","keywords":"Cloning (programming); clone (Java method); Code (set theory); Software maintenance; Commit; Programming language; Computer science; Source code; Software bug; Software evolution; Software; Biology; Software development; Database; Genetics; Software construction; Gene","score_opus":0.035743060351448065,"score_gpt":0.3307752259581475,"score_spread":0.2950321656066994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2972291021","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99732286,0.00015610803,0.0008817577,0.00011915201,0.0000058451,0.00005359862,0.000108578046,0.000019806484,0.0013323391],"genre_scores_gemma":[0.99816185,0.00011693053,0.00095805514,0.000048466194,0.0000066514003,0.0000358844,0.00017043985,0.000019165744,0.0004824928],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9911112,0.004288952,0.0010621066,0.00083790044,0.0022729344,0.0004267603],"domain_scores_gemma":[0.5099436,0.38553378,0.059505545,0.016574303,0.025294796,0.0031479585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009312757,0.00042585412,0.00030227663,0.0036408475,0.0013470299,0.0015064368,0.001401523,0.0012089775,0.0035620995],"category_scores_gemma":[0.16464168,0.0005145636,0.00040650685,0.0035272746,0.0019693184,0.0039385916,0.0011734779,0.0022250016,0.00050711766],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003254391,0.0016461943,0.96500653,0.00024396354,0.00009089913,0.00050546037,0.008037537,0.0003664933,0.0006890912,0.0009732565,0.0006136572,0.021501541],"study_design_scores_gemma":[0.000088711335,0.0016220235,0.96713924,0.00031739436,0.00026790774,0.0021205696,0.016943317,0.0058616833,0.0014386753,0.0010411935,0.0031079801,0.000051377105],"about_ca_topic_score_codex":0.0060185357,"about_ca_topic_score_gemma":0.007946117,"teacher_disagreement_score":0.009312757,"about_ca_system_score_codex":0.0010139868,"about_ca_system_score_gemma":0.0019544007,"threshold_uncertainty_score":0.0492512},"labels":[],"label_agreement":null},{"id":"W2972339742","doi":"10.26685/urncst.152","title":"Continuous Integration Using Gitlab","year":2019,"lang":"en","type":"article","venue":"Undergraduate Research in Natural and Clinical Science and Technology (URNCST) Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Agile software development; Computer science; Software engineering; Code (set theory); Software versioning; Quality (philosophy); Software; Engineering management; Programming language; Engineering","score_opus":0.053232967031571085,"score_gpt":0.40077815257360977,"score_spread":0.34754518554203867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2972339742","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04374931,0.0006910608,0.76454705,0.0010768978,0.00027130963,0.001065418,0.0012739144,0.16243228,0.024892826],"genre_scores_gemma":[0.18264961,0.00037113306,0.7859522,0.00063992006,0.00006276921,0.001322664,0.0036658957,0.014568111,0.010767672],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99416393,0.0016729725,0.0005481959,0.0010306811,0.0021076717,0.00047659254],"domain_scores_gemma":[0.9778717,0.00886968,0.0018215089,0.00669592,0.003952748,0.00078839075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066818637,0.0014601175,0.00068874325,0.0041637956,0.0010999228,0.0037145796,0.0030965586,0.0018786838,0.014559901],"category_scores_gemma":[0.03865409,0.001212645,0.0012248182,0.003370442,0.0017547107,0.00517593,0.008325631,0.003062875,0.006799744],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010688271,0.0005324649,0.013322614,0.0018990692,0.0002166132,0.0016601923,0.0133837145,0.013557882,0.024384577,0.02219723,0.05056236,0.85721445],"study_design_scores_gemma":[0.00087420514,0.0017091273,0.019239713,0.0021425271,0.0005200793,0.004235366,0.005232845,0.1962111,0.05301884,0.062224884,0.6537584,0.0008328918],"about_ca_topic_score_codex":0.003543591,"about_ca_topic_score_gemma":0.0031593898,"teacher_disagreement_score":0.014559901,"about_ca_system_score_codex":0.0009840853,"about_ca_system_score_gemma":0.0031066444,"threshold_uncertainty_score":0.048707724},"labels":[],"label_agreement":null},{"id":"W2973296283","doi":"10.1109/tse.2019.2941880","title":"Studying the Impact of Noises in Build Breakage Data","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Replicate; Breakage; Server; Troubleshooting; Timeout; Data mining; Data science; World Wide Web; Statistics; Operating system","score_opus":0.028169490543929115,"score_gpt":0.2874072299796261,"score_spread":0.25923773943569695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973296283","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92641056,0.0022963672,0.046055168,0.0014548409,0.0002719241,0.0002468513,0.019779481,0.0014928672,0.0019919882],"genre_scores_gemma":[0.9183797,0.0005034966,0.027362704,0.0004347799,0.00015422616,0.00021745767,0.05196543,0.00026876468,0.0007133158],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97417724,0.00901694,0.0027713676,0.0068251453,0.0060557583,0.0011535463],"domain_scores_gemma":[0.78962535,0.16123354,0.017718391,0.021536443,0.00810322,0.0017830126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02222134,0.0017119925,0.0011884323,0.005724118,0.0011744557,0.0030294706,0.0022314962,0.0025414096,0.00076428434],"category_scores_gemma":[0.120219305,0.0010696872,0.0022155938,0.0072757746,0.0022832735,0.0048682513,0.002670206,0.003545919,0.0006528511],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063003576,0.00049468956,0.8586049,0.00065742765,0.00071562914,0.0007393794,0.0024418714,0.08708007,0.0023728972,0.0027629368,0.00686986,0.036630183],"study_design_scores_gemma":[0.00011995005,0.00050338905,0.60936385,0.00047540216,0.0005280126,0.001469314,0.0034593863,0.33902267,0.007897339,0.010545905,0.02637993,0.00023473905],"about_ca_topic_score_codex":0.018862357,"about_ca_topic_score_gemma":0.020744963,"teacher_disagreement_score":0.02222134,"about_ca_system_score_codex":0.001913406,"about_ca_system_score_gemma":0.0018311989,"threshold_uncertainty_score":0.11751908},"labels":[],"label_agreement":null},{"id":"W2975484362","doi":"10.22215/etd/2019-13526","title":"Identifying Software Defects Using Neural Graph Classifiers","year":2019,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Suite; Java; Artificial intelligence; Graph; Source code; Software; Machine learning; Test suite; Task (project management); Software suite; Test case; Data mining; Programming language; Theoretical computer science; Engineering","score_opus":0.04384843707961259,"score_gpt":0.3148986019848136,"score_spread":0.271050164905201,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2975484362","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28121874,0.0022323634,0.69354236,0.0007740001,0.00024855527,0.00031797285,0.0014880733,0.0074470467,0.012730842],"genre_scores_gemma":[0.8517763,0.00071829115,0.1335193,0.0002643887,0.00010328808,0.00014312954,0.0037222952,0.0003078447,0.00944524],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9994273,0.00007314448,0.000027322216,0.00022697049,0.00015632276,0.00008901441],"domain_scores_gemma":[0.99842864,0.0007060469,0.00016783827,0.00016519208,0.00046686525,0.00006534915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00072031206,0.0009150195,0.00074898126,0.0035233332,0.00040614366,0.0011260391,0.0010745631,0.0013617675,0.0022506777],"category_scores_gemma":[0.0026021034,0.00029802747,0.0010372541,0.0012795348,0.00044925162,0.001126796,0.0004924895,0.00094715826,0.0012959309],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028255687,0.0003127767,0.009847772,0.00014160534,0.0001412098,0.00019992897,0.00007734538,0.1651853,0.01115603,0.0045071426,0.0118323425,0.796316],"study_design_scores_gemma":[0.000006718807,0.000049905964,0.0014267339,0.000020872862,0.000028683962,0.000044432432,0.000024911029,0.9908325,0.0025364822,0.0041963,0.00082604255,0.000006459171],"about_ca_topic_score_codex":0.009594243,"about_ca_topic_score_gemma":0.010393644,"teacher_disagreement_score":0.009594243,"about_ca_system_score_codex":0.0010245162,"about_ca_system_score_gemma":0.0006928295,"threshold_uncertainty_score":0.019076765},"labels":[],"label_agreement":null},{"id":"W2976928731","doi":"10.1016/j.jss.2019.110427","title":"An empirical study of security warnings from static application security testing tools","year":2019,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Waterloo","funders":"","keywords":"Computer science; False positive paradox; Source code; Application security; Vulnerability (computing); Computer security; Static program analysis; Static analysis; Code (set theory); Open source; Web application security; Software security assurance; Empirical research; Secure coding; Software; Data mining; Information security; Software development; The Internet; World Wide Web; Machine learning; Security service; Set (abstract data type); Programming language","score_opus":0.027078499168745673,"score_gpt":0.30272429123481737,"score_spread":0.2756457920660717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2976928731","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99859124,0.00004422081,0.000487126,0.0000619564,0.0000048767442,0.000035979836,0.000055973585,0.000011785139,0.00070688484],"genre_scores_gemma":[0.99905056,0.000030412037,0.0005203596,0.000025291572,0.0000034481454,0.000025668307,0.000109475455,0.0000083528685,0.00022653941],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99237424,0.003190115,0.00080630754,0.0005011481,0.0027873788,0.00034079968],"domain_scores_gemma":[0.5511046,0.37054953,0.04404118,0.010701861,0.020989658,0.0026131703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010056992,0.00044522248,0.00027855206,0.002441551,0.0005882729,0.0013173041,0.0012584464,0.0011288519,0.0026630762],"category_scores_gemma":[0.19409896,0.00036514626,0.00040902063,0.0016728546,0.0012358144,0.0031758514,0.0012994219,0.002282852,0.00049968244],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017139849,0.0076014586,0.92368656,0.00035401713,0.00014347241,0.0006070507,0.01161525,0.0013935363,0.0024530683,0.0016586676,0.0010484591,0.047724385],"study_design_scores_gemma":[0.00020882925,0.006400089,0.9524554,0.00027882363,0.00025603187,0.001193688,0.012742993,0.017729878,0.004432675,0.0013546313,0.002859852,0.000087124405],"about_ca_topic_score_codex":0.00223854,"about_ca_topic_score_gemma":0.0029981008,"teacher_disagreement_score":0.010056992,"about_ca_system_score_codex":0.0009121716,"about_ca_system_score_gemma":0.001211864,"threshold_uncertainty_score":0.053187072},"labels":[],"label_agreement":null},{"id":"W2978451778","doi":"10.82308/41899","title":"Dynamic data structure analysis and visualization of Java programs","year":2006,"lang":"en","type":"article","venue":"eScholarship@McGill (McGill)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Computer science; Java; Visualization; Data visualization; Computer graphics (images); Programming language; Data mining","score_opus":0.015926519627309756,"score_gpt":0.2651596914814227,"score_spread":0.24923317185411295,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2978451778","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1029783,0.000841583,0.8315864,0.0012495755,0.00017229142,0.00010316837,0.0024703296,0.04137877,0.019219525],"genre_scores_gemma":[0.49134964,0.0010792405,0.4916172,0.00021777774,0.00008394987,0.00018806281,0.0021138866,0.0062541487,0.007096063],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997292,0.000049423954,0.000014721105,0.00004594814,0.00012200158,0.000038559236],"domain_scores_gemma":[0.99905545,0.0003305835,0.00012357869,0.00013351266,0.00027719027,0.00007965527],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004593707,0.0005432633,0.00029715817,0.0023683514,0.0005081819,0.0016060262,0.0005669954,0.00038118925,0.00495192],"category_scores_gemma":[0.002005867,0.00028892801,0.0004449216,0.001156344,0.00040257428,0.0011109562,0.0010604332,0.0010187285,0.0007718598],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006145889,0.00024383639,0.014786936,0.00069850805,0.00009345211,0.0009935065,0.007379517,0.069087274,0.16799928,0.069534816,0.042077467,0.6264908],"study_design_scores_gemma":[0.000107539025,0.00015592866,0.020760028,0.000392895,0.00007640035,0.0009121319,0.0013479487,0.67146975,0.11521015,0.049048882,0.14031655,0.00020187986],"about_ca_topic_score_codex":0.0056130914,"about_ca_topic_score_gemma":0.0037183347,"teacher_disagreement_score":0.0056130914,"about_ca_system_score_codex":0.0006499633,"about_ca_system_score_gemma":0.0008740503,"threshold_uncertainty_score":0.01656586},"labels":[],"label_agreement":null},{"id":"W2979450652","doi":"10.1109/icmla.2019.00161","title":"A Multi-label, Dual-Output Deep Neural Network for Automated Bug Triaging","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Artificial intelligence; Lift (data mining); Chart; Process (computing); Interim; Artificial neural network; Machine learning; Deep learning; Scheme (mathematics); Programming language","score_opus":0.03329090948568393,"score_gpt":0.29875260057032577,"score_spread":0.26546169108464185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2979450652","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2052083,0.0016223934,0.7704463,0.0014806828,0.0004926599,0.00013677696,0.0011881219,0.01243446,0.0069902497],"genre_scores_gemma":[0.823055,0.0003293871,0.16387375,0.0006214629,0.000065553664,0.00012812526,0.0021474971,0.00013979114,0.009639359],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999597,0.0000670291,0.000022148733,0.00013361816,0.00009623017,0.00008390648],"domain_scores_gemma":[0.99924123,0.00021327594,0.000111487745,0.00010578434,0.0002603726,0.00006787251],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00096857693,0.0010088725,0.00050703593,0.00071602495,0.00034412346,0.00070288876,0.0020371694,0.0012695665,0.0015398118],"category_scores_gemma":[0.0026223888,0.00048579834,0.0005170642,0.0005873322,0.0004507964,0.0013601837,0.0011343138,0.0019300019,0.0007189976],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072212965,0.00073124137,0.008447343,0.000179128,0.0001657583,0.0003261092,0.00019853131,0.37078494,0.01925905,0.0031161495,0.017052852,0.5790168],"study_design_scores_gemma":[0.0000117504605,0.00004575451,0.00042891953,0.000011330578,0.000014699904,0.000020201289,0.00001137282,0.99461186,0.003151483,0.001122067,0.0005623201,0.000008193317],"about_ca_topic_score_codex":0.011578438,"about_ca_topic_score_gemma":0.016254997,"teacher_disagreement_score":0.011578438,"about_ca_system_score_codex":0.0013100547,"about_ca_system_score_gemma":0.0011524975,"threshold_uncertainty_score":0.023022115},"labels":[],"label_agreement":null},{"id":"W2981779826","doi":"10.1007/978-3-030-33702-5_5","title":"Towards Automated Microservices Extraction Using Muti-objective Evolutionary Search","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Microservices; Computer science; Benchmark (surveying); Sorting; Identification (biology); Source code; Set (abstract data type); Software; Software evolution; Plug-in; Genetic programming; Process (computing); Theoretical computer science; Data mining; Machine learning; Software development; Programming language; Cloud computing; Software construction; Operating system","score_opus":0.02359947469198001,"score_gpt":0.29985626001772736,"score_spread":0.2762567853257474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2981779826","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027241502,0.00039879777,0.96162724,0.000105299085,0.000054025688,0.0001284521,0.00036045923,0.004548647,0.005535642],"genre_scores_gemma":[0.13356608,0.0002585787,0.8558306,0.000069748196,0.00002046842,0.00011504689,0.0016837938,0.0007536345,0.0077019194],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994809,0.00007405697,0.000031294356,0.00012859094,0.00020883817,0.0000763022],"domain_scores_gemma":[0.9993679,0.00026186492,0.000060165952,0.00014975584,0.0001363086,0.00002410351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00048120803,0.0011375812,0.0011833187,0.0021212606,0.00075007306,0.001744159,0.0014763972,0.0014114744,0.0077795642],"category_scores_gemma":[0.0018609148,0.00069349445,0.0018398741,0.0019530281,0.00044947973,0.0013289725,0.0014599243,0.0010606878,0.0029745477],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016453276,0.00022487094,0.0018202435,0.00049164484,0.00011038867,0.0003358272,0.0001648563,0.14240971,0.043241832,0.014933758,0.0069935718,0.78910875],"study_design_scores_gemma":[0.00001965226,0.000060896848,0.0006781572,0.00004290197,0.00004256011,0.00018998161,0.0000949702,0.9617422,0.018809015,0.01158318,0.0067195864,0.000016960026],"about_ca_topic_score_codex":0.002521019,"about_ca_topic_score_gemma":0.004419444,"teacher_disagreement_score":0.0077795642,"about_ca_system_score_codex":0.00055229204,"about_ca_system_score_gemma":0.0011301156,"threshold_uncertainty_score":0.026025236},"labels":[],"label_agreement":null},{"id":"W2982011245","doi":"10.7202/1065133ar","title":"brat Rapid Annotation Tool. Web-based annotation and visualization tool","year":2019,"lang":"en","type":"article","venue":"Renaissance and Reformation","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Acadia University","funders":"","keywords":"Annotation; Computer science; Visualization; World Wide Web; Information retrieval; Artificial intelligence","score_opus":0.009773954082737747,"score_gpt":0.25010604564351196,"score_spread":0.24033209156077423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2982011245","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014154803,0.0008856392,0.27592832,0.0012288534,0.001392873,0.0010049611,0.19258273,0.49266773,0.03289349],"genre_scores_gemma":[0.012978828,0.0013835365,0.38080457,0.0018823452,0.00037477503,0.0068644243,0.34864116,0.16175705,0.08531343],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99611366,0.00077743107,0.0005095533,0.00053437636,0.0018030404,0.00026187638],"domain_scores_gemma":[0.98446333,0.007045632,0.00079744373,0.0017226624,0.0054463563,0.00052468455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006293354,0.0036606377,0.0025281035,0.011469662,0.0020818908,0.004838427,0.00415668,0.00304396,0.17588837],"category_scores_gemma":[0.01757038,0.0023964448,0.0016768322,0.0070067146,0.00088160654,0.006568038,0.0050980067,0.004317204,0.15836385],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032159622,0.00009288304,0.00038018936,0.0018045472,0.000056062163,0.0003443091,0.00058148865,0.00046537683,0.0065963953,0.0033432727,0.9206615,0.06535249],"study_design_scores_gemma":[0.00022721368,0.000050248596,0.0015069506,0.0008098738,0.000059859198,0.00048237017,0.0004023031,0.007328713,0.01665692,0.011385491,0.9608264,0.00026365244],"about_ca_topic_score_codex":0.0058540883,"about_ca_topic_score_gemma":0.008983319,"teacher_disagreement_score":0.17588837,"about_ca_system_score_codex":0.0012936174,"about_ca_system_score_gemma":0.0026216828,"threshold_uncertainty_score":0.58840525},"labels":[],"label_agreement":null},{"id":"W2982119098","doi":"10.1016/j.infsof.2019.106205","title":"Automatic prediction of the severity of bugs using stack traces and categorical features","year":2019,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Categorical variable; Computer science; Software bug; Software regression; Data mining; Machine learning; Software; Classifier (UML); Tracing; Predictive modelling; Artificial intelligence; Reliability engineering; Software quality; Software development; Engineering; Operating system","score_opus":0.007654808858986597,"score_gpt":0.22442895162658016,"score_spread":0.21677414276759358,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2982119098","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8345205,0.00081829535,0.14010294,0.00018745773,0.00010788831,0.00012474782,0.0056472747,0.017679509,0.0008114536],"genre_scores_gemma":[0.9373955,0.00015201271,0.05739589,0.000025851861,0.000037125814,0.000039508992,0.004200776,0.00018538792,0.0005679897],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999008,0.00011244354,0.00010801736,0.00026161538,0.0003992009,0.00011077556],"domain_scores_gemma":[0.98786956,0.0048705135,0.0028066533,0.0009710757,0.0026852703,0.00079691893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007598487,0.0014389305,0.0008667785,0.0068198517,0.00028242287,0.00085224095,0.0010017024,0.0010145685,0.001081706],"category_scores_gemma":[0.0071311085,0.0004257068,0.001018567,0.0021388237,0.00029292091,0.0014214518,0.0007721838,0.0008992935,0.0007033036],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017431115,0.0009099254,0.46655342,0.00076278683,0.0003994863,0.00091564626,0.00026481802,0.049164437,0.070725694,0.0016629903,0.008232333,0.39866546],"study_design_scores_gemma":[0.00008482918,0.00056454114,0.0899443,0.000057878595,0.00017823257,0.00047333227,0.00011330783,0.88637286,0.017710177,0.0031938858,0.0012331324,0.000073558185],"about_ca_topic_score_codex":0.005677462,"about_ca_topic_score_gemma":0.009116802,"teacher_disagreement_score":0.0068198517,"about_ca_system_score_codex":0.0003840643,"about_ca_system_score_gemma":0.00117365,"threshold_uncertainty_score":0.011288822},"labels":[],"label_agreement":null},{"id":"W2984281377","doi":"10.1145/3356773.3356811","title":"Continuous Data-driven Software Engineering - Towards a Research Agenda","year":2019,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software deployment; Pace; Software engineering; Software development; Software; Computer science; Social software engineering; Software Engineering Process Group; Automation; Systems engineering; Engineering; Engineering management; Software development process; Software construction","score_opus":0.07028905642559609,"score_gpt":0.32203960858503855,"score_spread":0.2517505521594425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2984281377","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02264136,0.26395544,0.45807987,0.22153944,0.0040243655,0.0004227104,0.0005695605,0.002080931,0.026686296],"genre_scores_gemma":[0.25407887,0.28156388,0.43480316,0.012731344,0.0041566826,0.0009024778,0.0018758735,0.0010150354,0.008872681],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9825873,0.008243322,0.0009516429,0.0020101583,0.0051206364,0.0010870352],"domain_scores_gemma":[0.83249944,0.12702727,0.0033345176,0.011451638,0.020554941,0.005132224],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04187352,0.0014511406,0.001994541,0.004425041,0.0013330602,0.01443873,0.0054818997,0.008493579,0.004732392],"category_scores_gemma":[0.06359564,0.0014383716,0.0014676767,0.007614761,0.0076342546,0.0363348,0.0062625175,0.009697287,0.0017392021],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026349595,0.00087581686,0.0034927672,0.004495943,0.000089421184,0.000182141,0.0022163314,0.019050026,0.0020130852,0.54157984,0.022072205,0.40366888],"study_design_scores_gemma":[0.00013129112,0.00048806678,0.0012511835,0.005383927,0.000087236585,0.00025054807,0.003958357,0.06866845,0.003959014,0.7004292,0.21520561,0.00018705096],"about_ca_topic_score_codex":0.0052248524,"about_ca_topic_score_gemma":0.0017520548,"teacher_disagreement_score":0.9581265,"about_ca_system_score_codex":0.005946693,"about_ca_system_score_gemma":0.0136206895,"threshold_uncertainty_score":0.22145098},"labels":[],"label_agreement":null},{"id":"W2985253063","doi":"10.1145/3356773.3356806","title":"SST'19 - Software and Systems Traceability","year":2019,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Traceability; Requirements traceability; Software engineering; Computer science; Context (archaeology); Software development; Software system; Change impact analysis; Event (particle physics); Systems engineering; Software; Engineering; Requirement; Programming language","score_opus":0.016684252597717532,"score_gpt":0.24249219578245296,"score_spread":0.22580794318473543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2985253063","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0085920375,0.008369808,0.4969756,0.041007426,0.0232322,0.0024928255,0.0063646575,0.013430198,0.39953527],"genre_scores_gemma":[0.1267687,0.023267154,0.357054,0.015780445,0.015935786,0.0031000718,0.043153446,0.012928613,0.40201178],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9648586,0.0075705056,0.0037551045,0.0028603955,0.018609941,0.002345512],"domain_scores_gemma":[0.964112,0.0076824217,0.0014595925,0.0056158486,0.017949758,0.0031804012],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02475603,0.0030402015,0.002537191,0.0049884384,0.0038948893,0.011059213,0.004603861,0.007831379,0.031892885],"category_scores_gemma":[0.03079936,0.0016149092,0.0030868836,0.0049306783,0.0056376434,0.009639454,0.007804585,0.0090515725,0.025165819],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043747635,0.00038493404,0.0015210143,0.0013629422,0.00014770414,0.0009284452,0.0012292395,0.011200637,0.00984062,0.29816312,0.42236465,0.25241923],"study_design_scores_gemma":[0.0001011697,0.00031219693,0.0008997392,0.0006853117,0.000042034593,0.00062569306,0.00022279198,0.00555739,0.006786962,0.073145024,0.911528,0.000093722876],"about_ca_topic_score_codex":0.017466512,"about_ca_topic_score_gemma":0.012015397,"teacher_disagreement_score":0.031892885,"about_ca_system_score_codex":0.0058037317,"about_ca_system_score_gemma":0.0148701295,"threshold_uncertainty_score":0.13092399},"labels":[],"label_agreement":null},{"id":"W2987292856","doi":"10.1016/j.infsof.2019.106213","title":"Towards assisting developers in API usage by automated recovery of complex temporal patterns","year":2019,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Rimouski","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Application programming interface; Software design pattern; Usage data; Software; Programming language; World Wide Web","score_opus":0.012240438639937762,"score_gpt":0.24952561703720383,"score_spread":0.23728517839726607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2987292856","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06941648,0.0003686087,0.88023853,0.0017933869,0.00009597839,0.00038953745,0.0007652652,0.043997716,0.002934538],"genre_scores_gemma":[0.2773747,0.00028111116,0.71505904,0.0003547519,0.000042622036,0.00012829015,0.0015357066,0.0016702343,0.0035534545],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99369663,0.0022678846,0.00050251436,0.0010963792,0.0020561377,0.00038044644],"domain_scores_gemma":[0.9665451,0.011294269,0.004912611,0.010633444,0.0057759886,0.0008386115],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054106787,0.0015131781,0.00086316315,0.002592969,0.0007130703,0.0036484317,0.002736333,0.0022270214,0.0027679354],"category_scores_gemma":[0.033802833,0.0009690093,0.00091427256,0.0012626703,0.0007263667,0.003962398,0.0031710379,0.0027497069,0.0035357536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054652157,0.0012080502,0.032006063,0.00076700805,0.00015213525,0.000713823,0.0027475473,0.019444698,0.081468634,0.005054472,0.022853673,0.8330374],"study_design_scores_gemma":[0.00011934666,0.00035853425,0.011033342,0.00033887182,0.000182166,0.0010948278,0.0017431997,0.83111876,0.10537018,0.016882407,0.031599008,0.00015924676],"about_ca_topic_score_codex":0.004626368,"about_ca_topic_score_gemma":0.008995463,"teacher_disagreement_score":0.0054106787,"about_ca_system_score_codex":0.0005036583,"about_ca_system_score_gemma":0.0033954051,"threshold_uncertainty_score":0.02861476},"labels":[],"label_agreement":null},{"id":"W2987358222","doi":"10.1109/vissoft.2019.00013","title":"A Tertiary Systematic Literature Review on Software Visualization","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Visualization; Software visualization; Computer science; Scope (computer science); Software; Data science; Systematic review; Software engineering; Field (mathematics); Disconnection; Software development; Human–computer interaction; Software construction; Data mining","score_opus":0.009328171718587968,"score_gpt":0.27461267983293786,"score_spread":0.2652845081143499,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2987358222","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004180841,0.9863988,0.002354111,0.0016432987,0.0005193352,0.0010534904,0.0017753949,0.00006245261,0.002012257],"genre_scores_gemma":[0.027966667,0.959164,0.005782582,0.0021380847,0.00026053438,0.002404365,0.0016595782,0.000044239656,0.0005799094],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9821745,0.0065824743,0.0063110683,0.0012838561,0.0031533246,0.0004947354],"domain_scores_gemma":[0.89644945,0.07262649,0.00932182,0.0027427883,0.01769843,0.0011610484],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.020776212,0.0015202906,0.0044705695,0.036120713,0.0014753992,0.0039814585,0.0020134125,0.002317556,0.0073203626],"category_scores_gemma":[0.09812531,0.0010838427,0.005036331,0.020830462,0.001456837,0.004651479,0.0033557876,0.0016332663,0.0011371353],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018076,0.00005152519,0.0014085505,0.839445,0.0020395378,0.00021675338,0.0018317167,0.00015357984,0.0010043406,0.0016464335,0.008223861,0.14379796],"study_design_scores_gemma":[0.0000716303,0.00017052826,0.0026786537,0.9394351,0.008247459,0.00031879434,0.001183238,0.000085488085,0.00047462204,0.0009670735,0.046333056,0.000034436634],"about_ca_topic_score_codex":0.0058380812,"about_ca_topic_score_gemma":0.018439505,"teacher_disagreement_score":0.9792238,"about_ca_system_score_codex":0.0052294023,"about_ca_system_score_gemma":0.02939262,"threshold_uncertainty_score":0.109876454},"labels":[],"label_agreement":null},{"id":"W2988126082","doi":"10.1109/vissoft.2019.00019","title":"CloneCompass: Visualizations for Exploring Assembly Code Clone Ecosystems","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; clone (Java method); Source code; Code (set theory); Codebase; Software engineering; Software; Function (biology); Visualization; Process (computing); Programming language; Data mining","score_opus":0.07011543209215279,"score_gpt":0.31788373295285205,"score_spread":0.24776830086069926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2988126082","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.116494484,0.00077574735,0.7808136,0.0013192862,0.00016487436,0.0005440567,0.009362656,0.07811169,0.012413568],"genre_scores_gemma":[0.32256353,0.0006003719,0.66111654,0.00022327303,0.000051890827,0.00086642976,0.0046529127,0.0066239885,0.003301076],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9994863,0.0001729718,0.00004177037,0.00008585433,0.00017028183,0.00004279095],"domain_scores_gemma":[0.99319756,0.0048788083,0.00038568935,0.00057251117,0.00071720633,0.00024815436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017835042,0.001191263,0.0004277289,0.0037912955,0.00069872965,0.0022024466,0.0010538619,0.001008774,0.013381136],"category_scores_gemma":[0.010317825,0.0005774937,0.0007246362,0.0016588231,0.0004990829,0.0027485017,0.0028141928,0.0009921234,0.001218888],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023607882,0.00062134716,0.025717467,0.0037622924,0.00032196537,0.0029377819,0.038392037,0.041861944,0.11018198,0.062476348,0.11489635,0.59646976],"study_design_scores_gemma":[0.000562322,0.0007880547,0.028210294,0.0012191648,0.00022307783,0.0025704633,0.0076219565,0.52857196,0.09088943,0.06804139,0.27078828,0.000513607],"about_ca_topic_score_codex":0.0029322042,"about_ca_topic_score_gemma":0.005076994,"teacher_disagreement_score":0.013381136,"about_ca_system_score_codex":0.00056802965,"about_ca_system_score_gemma":0.0008179246,"threshold_uncertainty_score":0.04476434},"labels":[],"label_agreement":null},{"id":"W2989443621","doi":"10.1109/tse.2019.2952130","title":"An Empirical Study of Dependency Downgrades in the npm Ecosystem","year":2019,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Downgrade; Dependency (UML); Software versioning; Reuse; Software; Software engineering; Computer security; Operating system","score_opus":0.01786606815314012,"score_gpt":0.28024781358299794,"score_spread":0.2623817454298578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2989443621","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9949597,0.00023213061,0.0012926115,0.0004376889,0.000010617639,0.00005878188,0.00026416717,0.00005192393,0.002692236],"genre_scores_gemma":[0.99719954,0.00017606006,0.0014416801,0.00009295469,0.000014269019,0.000048794573,0.0005006304,0.00003750519,0.0004885683],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9886565,0.004677258,0.0013342127,0.0014916646,0.003129233,0.00071114185],"domain_scores_gemma":[0.6695696,0.19338502,0.08991594,0.01883767,0.022716327,0.0055754716],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015534879,0.0003573726,0.0003171775,0.0029444417,0.0014929866,0.0029682384,0.0014983389,0.001106839,0.0029059625],"category_scores_gemma":[0.15066732,0.00058486365,0.0004217536,0.0033948438,0.0021856835,0.007066883,0.0022975479,0.0026265027,0.00056177744],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014225743,0.00038891967,0.9388191,0.0002576807,0.000059445094,0.00069094996,0.028902097,0.0004645716,0.0006988836,0.0017988522,0.0017718843,0.026005333],"study_design_scores_gemma":[0.000018982944,0.00023400123,0.9516799,0.00022972815,0.000052055922,0.0007951373,0.030175764,0.0042097755,0.0007516078,0.0016141052,0.01018018,0.000058847065],"about_ca_topic_score_codex":0.006871186,"about_ca_topic_score_gemma":0.0066953213,"teacher_disagreement_score":0.9844651,"about_ca_system_score_codex":0.0019994036,"about_ca_system_score_gemma":0.0014356675,"threshold_uncertainty_score":0.082157254},"labels":[],"label_agreement":null},{"id":"W2989967080","doi":"","title":"Code Forking and Software Development Project Sustainability: Evidence from GitHub","year":2019,"lang":"en","type":"article","venue":"Journal of the Association for Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal","funders":"","keywords":"Computer science; Software engineering; Software development; Code (set theory); Software; Sustainability; Programming language; Ecology","score_opus":0.021934243268333514,"score_gpt":0.273981560421965,"score_spread":0.2520473171536315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2989967080","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99569595,0.0006243825,0.00046281947,0.00058679114,0.0000062220183,0.000025457803,0.00031213468,0.000021634725,0.002264628],"genre_scores_gemma":[0.99794596,0.00047478653,0.00042358285,0.00009392775,0.000010200142,0.000043585962,0.00041692198,0.000031121155,0.0005599044],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99220663,0.0034070264,0.0005310574,0.0012805297,0.001744519,0.00083017023],"domain_scores_gemma":[0.83923805,0.06852365,0.06086047,0.010283239,0.013523852,0.0075707054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010755011,0.00038987398,0.0004190389,0.0049942974,0.0018263671,0.0026074143,0.0015839323,0.0010668031,0.003353001],"category_scores_gemma":[0.08799549,0.00047074724,0.0004346423,0.0081604235,0.0038060953,0.0050389622,0.0055601983,0.0016066533,0.00062573777],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015642899,0.00022040844,0.9544202,0.00019311225,0.00009441115,0.000217263,0.02074795,0.00015781292,0.00014216246,0.00086507585,0.00096335297,0.021821922],"study_design_scores_gemma":[0.000012746934,0.000098143464,0.9816539,0.00022797853,0.000045550503,0.00014266084,0.013588764,0.0006452437,0.00016671092,0.0010727015,0.0023221,0.000023430603],"about_ca_topic_score_codex":0.033875912,"about_ca_topic_score_gemma":0.05347942,"teacher_disagreement_score":0.033875912,"about_ca_system_score_codex":0.002197415,"about_ca_system_score_gemma":0.0027473215,"threshold_uncertainty_score":0.06735748},"labels":[],"label_agreement":null},{"id":"W2990774762","doi":"10.1109/models.2019.00-19","title":"Pitfalls Analyzer: Quality Control for Model-Driven Data Science Pipelines","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Pipeline transport; Pipeline (software); Data science; Raw data; Quality (philosophy); Data mining; Software engineering; Programming language; Engineering","score_opus":0.07607110023299138,"score_gpt":0.3725011453655077,"score_spread":0.29643004513251636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2990774762","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016025348,0.0002696137,0.681366,0.00041613783,0.00011146795,0.00054695667,0.0021933354,0.29796377,0.0011073481],"genre_scores_gemma":[0.24075824,0.0003372345,0.7225288,0.0006239291,0.00007796269,0.00095756655,0.008706591,0.02420336,0.0018063494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.982226,0.0035142421,0.0023603248,0.00314067,0.0078179315,0.00094083627],"domain_scores_gemma":[0.92171586,0.03775695,0.010030228,0.016546736,0.0124900825,0.001460116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018064437,0.0035669822,0.0013780796,0.0059724525,0.0011698857,0.004554213,0.0054116515,0.0016285806,0.0055735684],"category_scores_gemma":[0.07953137,0.0027809315,0.002749533,0.0025925091,0.0024878003,0.007918619,0.005213135,0.004282668,0.0019720888],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038987568,0.001163546,0.07033359,0.0046248706,0.0010516434,0.0020196724,0.004309896,0.11928153,0.07983488,0.04533291,0.1127512,0.55539757],"study_design_scores_gemma":[0.0003203163,0.00042903505,0.007796142,0.00035920454,0.00021006746,0.00059637375,0.00033302463,0.8260205,0.09904215,0.024274906,0.040285196,0.00033308784],"about_ca_topic_score_codex":0.010154874,"about_ca_topic_score_gemma":0.007580103,"teacher_disagreement_score":0.018064437,"about_ca_system_score_codex":0.0029153728,"about_ca_system_score_gemma":0.0068421215,"threshold_uncertainty_score":0.09553504},"labels":[],"label_agreement":null},{"id":"W2993280551","doi":"10.1109/icsme.2019.00026","title":"Investigating Context Adaptation Bugs in Code Clones","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Cloning (programming); Context (archaeology); Computer science; clone (Java method); Code (set theory); Programming language; Fragment (logic); Source code; Java; Software bug; Software maintenance; Biology; Software; Software system; Genetics; Gene","score_opus":0.035022700386172355,"score_gpt":0.26918638135110995,"score_spread":0.2341636809649376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2993280551","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9710706,0.0010233119,0.025584463,0.000090208734,0.000027424696,0.00014277323,0.0003368584,0.000749781,0.00097444915],"genre_scores_gemma":[0.9717605,0.00032255388,0.026419148,0.000063129766,0.00001999379,0.000092771654,0.00060844095,0.00014103437,0.00057233806],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99552417,0.0006801861,0.00039223544,0.0012663634,0.0018644362,0.00027264157],"domain_scores_gemma":[0.9379358,0.033539545,0.016615424,0.004446554,0.0066465084,0.0008162121],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020959326,0.0005459249,0.0005109288,0.00429729,0.0006248343,0.001167773,0.0008118578,0.000838332,0.0006387603],"category_scores_gemma":[0.039625388,0.00037557588,0.00059176446,0.0026741403,0.00091425696,0.0022771866,0.0012407558,0.0007409412,0.00016833127],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003361817,0.00020323036,0.7891029,0.00090978306,0.00025888789,0.003290179,0.0060015144,0.0058420235,0.03145335,0.0029556244,0.0013757406,0.15827061],"study_design_scores_gemma":[0.00007154697,0.00076919986,0.81907314,0.00037567315,0.00064039556,0.008822251,0.004693607,0.10900755,0.038709138,0.006054234,0.011589068,0.00019425924],"about_ca_topic_score_codex":0.0033231385,"about_ca_topic_score_gemma":0.0039323187,"teacher_disagreement_score":0.00429729,"about_ca_system_score_codex":0.00065542426,"about_ca_system_score_gemma":0.0007839306,"threshold_uncertainty_score":0.011084497},"labels":[],"label_agreement":null},{"id":"W2993429484","doi":"10.1109/icsme.2019.00031","title":"Aiding Code Change Understanding with Semantic Change Impact Analysis","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Change impact analysis; Computer science; Semantic change; Context (archaeology); JavaScript; Unix; Code (set theory); False positive paradox; Source code; Software engineering; Programming language; Information retrieval; Data science; Artificial intelligence; Software","score_opus":0.11918472573368231,"score_gpt":0.31640870004469995,"score_spread":0.19722397431101762,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2993429484","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1363921,0.0008737128,0.8197194,0.0013831021,0.00016673416,0.00089751184,0.0023969808,0.032059204,0.0061112484],"genre_scores_gemma":[0.45519432,0.000409421,0.5372997,0.00026459375,0.000111288704,0.0003504579,0.0033998627,0.0014898928,0.0014804823],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9906579,0.001987015,0.0009129905,0.0014843119,0.004577827,0.00037993942],"domain_scores_gemma":[0.94307375,0.031303383,0.0082784165,0.0065933065,0.010333993,0.0004171891],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055905557,0.0017107364,0.0012009433,0.014003495,0.0010708428,0.00347105,0.0016595206,0.0013213973,0.0020967296],"category_scores_gemma":[0.04494857,0.0008081968,0.0016579515,0.005683126,0.0014289764,0.0052721086,0.0028827512,0.0027986849,0.0007951967],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063815847,0.00061005546,0.09713649,0.0015975598,0.00034843566,0.001312274,0.006713248,0.022211753,0.04664514,0.020705028,0.015625836,0.7864561],"study_design_scores_gemma":[0.00012434133,0.0005207112,0.07036316,0.0005589981,0.00054648967,0.0017524272,0.0033138325,0.7118118,0.09808985,0.06829092,0.044252656,0.0003748389],"about_ca_topic_score_codex":0.004054644,"about_ca_topic_score_gemma":0.006072884,"teacher_disagreement_score":0.014003495,"about_ca_system_score_codex":0.0016271028,"about_ca_system_score_gemma":0.00250775,"threshold_uncertainty_score":0.02956605},"labels":[],"label_agreement":null},{"id":"W2993561384","doi":"10.1109/re.2019.00055","title":"Data-Driven Elicitation and Optimization of Dependencies between Requirements","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Requirements elicitation; Data mining; Programming language; Requirements analysis; Software","score_opus":0.06475097418202472,"score_gpt":0.3167202874128103,"score_spread":0.25196931323078553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2993561384","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16596617,0.00025907805,0.8234533,0.00074054336,0.000033699187,0.00084466144,0.0034588834,0.0027332238,0.0025103965],"genre_scores_gemma":[0.3757462,0.00014031629,0.6128984,0.00013597406,0.000022461098,0.000723098,0.009042311,0.00019610614,0.0010952002],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995626,0.0017505191,0.0003745544,0.00085085013,0.0012245695,0.0001734044],"domain_scores_gemma":[0.9704936,0.022265632,0.0021069255,0.0016745611,0.0031889288,0.00027026108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042376337,0.0010309794,0.0007709313,0.002486254,0.0005431918,0.0011013226,0.0012219178,0.0008952159,0.0014598221],"category_scores_gemma":[0.022626301,0.00061742397,0.0011075437,0.0013347848,0.0006146336,0.0021464252,0.0013721569,0.001662462,0.00070700294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005338862,0.0011609531,0.025036223,0.001975864,0.00011812885,0.0006748496,0.001345074,0.22300501,0.04624048,0.00688115,0.0085433675,0.68448514],"study_design_scores_gemma":[0.00003692073,0.00017434487,0.00676211,0.000093195136,0.000037373895,0.00017500205,0.00041875706,0.9519999,0.025364745,0.009334635,0.0055608614,0.0000421627],"about_ca_topic_score_codex":0.0027932501,"about_ca_topic_score_gemma":0.0075867814,"teacher_disagreement_score":0.0042376337,"about_ca_system_score_codex":0.001080502,"about_ca_system_score_gemma":0.0031572971,"threshold_uncertainty_score":0.022411048},"labels":[],"label_agreement":null},{"id":"W2994133228","doi":"10.5555/1858449.1858458","title":"Using mock object frameworks to teach object-oriented design principles","year":2010,"lang":"en","type":"article","venue":"Journal of computing sciences in colleges","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Object-oriented design; Computer science; Object (grammar); Object-oriented programming; Method; Class (philosophy); Software engineering; Context (archaeology); Inheritance (genetic algorithm); Programming language; Artificial intelligence","score_opus":0.04980164223800232,"score_gpt":0.3451517191738768,"score_spread":0.2953500769358745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2994133228","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026601398,0.0010966854,0.9360276,0.0028256632,0.00032997297,0.00031736353,0.000028166432,0.002005891,0.030767301],"genre_scores_gemma":[0.11621955,0.0013299609,0.86406195,0.00091526436,0.00008136423,0.00043882077,0.00009898999,0.00033501108,0.016519101],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973948,0.0012716502,0.00012803833,0.00020369835,0.00083499117,0.00016680008],"domain_scores_gemma":[0.99569976,0.0021529377,0.00028318728,0.00087001355,0.0006006492,0.0003935256],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003844682,0.0010512422,0.00045711373,0.0011647297,0.0013902765,0.002772008,0.00163881,0.0017520598,0.0039918367],"category_scores_gemma":[0.010960083,0.0006442088,0.00069254666,0.000508049,0.0025979301,0.0030520048,0.0030185562,0.0022485876,0.0016480774],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000060319944,0.00047826325,0.0022383279,0.0005833948,0.000026719123,0.0006112627,0.01200627,0.014514364,0.013140708,0.6847243,0.019623913,0.25199214],"study_design_scores_gemma":[0.00008199658,0.00038660722,0.0013071242,0.0005794188,0.00003204137,0.0017898038,0.0018074688,0.02987009,0.017323406,0.39627543,0.55040044,0.00014614154],"about_ca_topic_score_codex":0.0007084575,"about_ca_topic_score_gemma":0.0019417874,"teacher_disagreement_score":0.0039918367,"about_ca_system_score_codex":0.0013010801,"about_ca_system_score_gemma":0.002152015,"threshold_uncertainty_score":0.020332873},"labels":[],"label_agreement":null},{"id":"W2994598966","doi":"10.1109/icsme.2019.00018","title":"Improving Bug Triaging with High Confidence Predictions at Ericsson","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Triage; Software bug; Categorical variable; Context (archaeology); Classifier (UML); Precision and recall; Replicate; Software regression; Machine learning; Artificial intelligence; Data mining; Software; Software quality; Statistics; Software development","score_opus":0.00884423578975883,"score_gpt":0.22233329105635766,"score_spread":0.21348905526659884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2994598966","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8630796,0.0037615253,0.088604875,0.004098343,0.00032560507,0.00013933482,0.0027097515,0.033478774,0.0038021556],"genre_scores_gemma":[0.93873435,0.0005096562,0.0525467,0.00041634767,0.0001343644,0.000053398013,0.0045129317,0.0006413102,0.0024508473],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99243426,0.0023274121,0.0006103006,0.0020021112,0.0020629542,0.00056297035],"domain_scores_gemma":[0.9507777,0.023889909,0.009632101,0.00546246,0.008896271,0.001341564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012555391,0.0019062681,0.0019209527,0.005681797,0.0007083744,0.003284124,0.0019668823,0.002213915,0.0014269882],"category_scores_gemma":[0.051973462,0.00095140183,0.0010854019,0.0032535999,0.00051605236,0.002886176,0.002031657,0.0028406626,0.003469815],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011752837,0.0012107758,0.4172883,0.0005488711,0.0006380577,0.0009933665,0.0013263598,0.16017596,0.007477721,0.0011926049,0.04718669,0.360786],"study_design_scores_gemma":[0.000090699796,0.00046536452,0.049116515,0.00014823469,0.00019664627,0.0006745076,0.00036246912,0.93306077,0.0077369073,0.0018667337,0.0061653373,0.00011580388],"about_ca_topic_score_codex":0.018865123,"about_ca_topic_score_gemma":0.017783646,"teacher_disagreement_score":0.018865123,"about_ca_system_score_codex":0.0015311606,"about_ca_system_score_gemma":0.0021938495,"threshold_uncertainty_score":0.06640005},"labels":[],"label_agreement":null},{"id":"W2994650698","doi":"10.1109/scam.2019.00016","title":"An Exploratory Study on Automatic Architectural Change Analysis Using Natural Language Processing Techniques","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Commit; Computer science; Software engineering; Software architecture; Categorization; Architecture; Outcome (game theory); Change impact analysis; Software; Artificial intelligence; Programming language; Database","score_opus":0.03385155722738631,"score_gpt":0.3349857242929544,"score_spread":0.3011341670655681,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2994650698","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9771573,0.00026579344,0.019173196,0.00024677778,0.000015359992,0.00072752737,0.0006180896,0.00021521487,0.0015808545],"genre_scores_gemma":[0.9271461,0.00027762793,0.068008274,0.00029339214,0.000039440616,0.00084851513,0.002140203,0.000105420695,0.0011410818],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9810657,0.012739856,0.00117026,0.0022731095,0.0023014983,0.00044962417],"domain_scores_gemma":[0.8151671,0.16293488,0.0039352607,0.0066414583,0.010558043,0.00076326926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014535191,0.00073059666,0.0006549525,0.0028120137,0.001113742,0.0016880037,0.0016104318,0.0011585094,0.001239323],"category_scores_gemma":[0.069957756,0.00036554158,0.00073150406,0.0024782962,0.0012414415,0.003512521,0.0012874906,0.0012585128,0.00068519614],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024463825,0.01610242,0.27779827,0.005305571,0.00030568335,0.0035035945,0.08418121,0.009049514,0.063834,0.0042260475,0.006891122,0.5263561],"study_design_scores_gemma":[0.0010574441,0.016511178,0.4743174,0.0013282167,0.0005801935,0.007736969,0.082980424,0.23011222,0.11702654,0.008860315,0.058967724,0.00052134495],"about_ca_topic_score_codex":0.0021595838,"about_ca_topic_score_gemma":0.0033857343,"teacher_disagreement_score":0.014535191,"about_ca_system_score_codex":0.0009702501,"about_ca_system_score_gemma":0.0010601712,"threshold_uncertainty_score":0.07687032},"labels":[],"label_agreement":null},{"id":"W2995099744","doi":"10.1109/rew.2019.00022","title":"What Can the Sentiment of a Software Requirements Specification Document Tell Us?","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Sentiment analysis; Software; Set (abstract data type); Data science; Information retrieval; Software engineering; World Wide Web; Natural language processing; Programming language","score_opus":0.022398692669749573,"score_gpt":0.27583511503056846,"score_spread":0.25343642236081887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2995099744","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94346267,0.001705298,0.011020053,0.006726579,0.0006749494,0.00019651037,0.005466545,0.00021049657,0.03053693],"genre_scores_gemma":[0.9871644,0.0009828894,0.0060130954,0.0008006893,0.00025983792,0.00011949678,0.0026395684,0.00008027134,0.0019397434],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99682164,0.0011626625,0.0003620991,0.00027715482,0.0011903337,0.00018617128],"domain_scores_gemma":[0.97504413,0.014147415,0.003648689,0.0004953269,0.0061263596,0.00053808617],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00400668,0.00037569692,0.0005114899,0.0023760858,0.00063081644,0.0030292445,0.00026394427,0.00065604044,0.0017919171],"category_scores_gemma":[0.029263685,0.00015912387,0.00040507782,0.0018380939,0.0007710547,0.00251095,0.0006203806,0.0008545886,0.00095570873],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024871004,0.00042368437,0.3253614,0.004132182,0.00058856583,0.0013682119,0.039563198,0.0017605834,0.06500881,0.009985186,0.065377325,0.4839438],"study_design_scores_gemma":[0.0000858489,0.0011662819,0.7420003,0.0011088899,0.0005018448,0.0016261137,0.06865653,0.023904502,0.027148366,0.013357314,0.120004274,0.0004397394],"about_ca_topic_score_codex":0.002146991,"about_ca_topic_score_gemma":0.0023440924,"teacher_disagreement_score":0.00400668,"about_ca_system_score_codex":0.0010705562,"about_ca_system_score_gemma":0.00051028735,"threshold_uncertainty_score":0.02118963},"labels":[],"label_agreement":null},{"id":"W2996318555","doi":"10.22215/etd/2014-10501","title":"Modeling Use-Case Sequential Dependencies Using ACL","year":2014,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Semantics (computer science); Programming language; Independence (probability theory); Ask price; Unary operation; Theoretical computer science; Mathematics","score_opus":0.0783807356310412,"score_gpt":0.326196916639899,"score_spread":0.24781618100885783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996318555","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014119335,0.00007744147,0.9771538,0.00029721082,0.000018828148,0.0003539355,0.0005476596,0.0031707685,0.0042610443],"genre_scores_gemma":[0.19760998,0.00028886992,0.79433644,0.00019480356,0.000043583772,0.0009521233,0.0019974394,0.00085075025,0.0037260551],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.992191,0.0028704475,0.0011644837,0.0007357398,0.002540184,0.0004981632],"domain_scores_gemma":[0.9752205,0.015099892,0.0023700523,0.0038482,0.0031396202,0.00032170655],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071778097,0.0013042178,0.00059930846,0.0038943565,0.000973375,0.0039357063,0.0025139553,0.0014871873,0.004972193],"category_scores_gemma":[0.02183187,0.0016343932,0.0020179795,0.0017612735,0.001927949,0.00617631,0.0024523176,0.002646119,0.0015775347],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029500687,0.0004811668,0.014346858,0.00086356903,0.00015794903,0.0022452304,0.003113765,0.25256738,0.014696247,0.51888746,0.008670209,0.18367508],"study_design_scores_gemma":[0.00009013519,0.0000834572,0.0009139406,0.00019343443,0.000114110764,0.0006540266,0.00026767625,0.78970504,0.017299019,0.13836946,0.05223063,0.000079085046],"about_ca_topic_score_codex":0.0100669805,"about_ca_topic_score_gemma":0.010745021,"teacher_disagreement_score":0.0100669805,"about_ca_system_score_codex":0.0016657709,"about_ca_system_score_gemma":0.004182814,"threshold_uncertainty_score":0.03796035},"labels":[],"label_agreement":null},{"id":"W2996771664","doi":"10.1109/saner48275.2020.9054792","title":"Cross-Dataset Design Discussion Mining","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Classifier (UML); Transfer of learning; Artificial intelligence; Machine learning; Documentation; Software; Replicate; Data mining","score_opus":0.08501070002276262,"score_gpt":0.35010330652625204,"score_spread":0.2650926065034894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996771664","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2972983,0.0070659653,0.12372713,0.0032338623,0.0011071595,0.0068948083,0.45750707,0.011930562,0.09123517],"genre_scores_gemma":[0.259554,0.0008047222,0.13176084,0.0012371012,0.00025559706,0.007140467,0.5757498,0.0010909869,0.02240647],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98470867,0.0050932383,0.0019093815,0.003558557,0.0040301303,0.00070005486],"domain_scores_gemma":[0.9562793,0.017694043,0.0039317994,0.012902575,0.008007907,0.0011845184],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009887331,0.0012867968,0.00074140733,0.011270844,0.0023328476,0.003417261,0.003113309,0.002401001,0.011431787],"category_scores_gemma":[0.04082252,0.0005274744,0.0015588313,0.008328183,0.0007764983,0.0028731986,0.0039182073,0.0021556849,0.0069211447],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012415503,0.0023071072,0.10111828,0.0063884766,0.00094637636,0.0009844906,0.0038084085,0.0070402836,0.013877151,0.015640942,0.3336844,0.5129625],"study_design_scores_gemma":[0.00035489834,0.0004626002,0.08795677,0.0007677053,0.0004051254,0.0010085221,0.003438568,0.025582694,0.023901427,0.012829558,0.84311366,0.00017848423],"about_ca_topic_score_codex":0.0035154065,"about_ca_topic_score_gemma":0.0070150997,"teacher_disagreement_score":0.011431787,"about_ca_system_score_codex":0.0018609228,"about_ca_system_score_gemma":0.0026599308,"threshold_uncertainty_score":0.052289844},"labels":[],"label_agreement":null},{"id":"W2997090216","doi":"10.1609/aaai.v34i04.5733","title":"Predicting Propositional Satisfiability via End-to-End Learning","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"British Columbia Knowledge Development Fund; Natural Sciences and Engineering Research Council of Canada; Alberta Machine Intelligence Institute; Compute Canada; Defense Advanced Research Projects Agency; Canadian Institute for Advanced Research; Nvidia","keywords":"Satisfiability; Computer science; Variable (mathematics); Boolean satisfiability problem; Feature (linguistics); Artificial intelligence; Deep learning; Feature engineering; Linear programming; Key (lock); Algorithm; Machine learning; Theoretical computer science; Mathematics","score_opus":0.06095954478794397,"score_gpt":0.2915961234970659,"score_spread":0.2306365787091219,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997090216","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16206945,0.0009127571,0.7927228,0.0027962457,0.00033759684,0.00041882868,0.003352469,0.024713794,0.0126759745],"genre_scores_gemma":[0.64947367,0.0004469775,0.33067772,0.0013768715,0.0001532043,0.00050430314,0.009696001,0.00086425274,0.006807057],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985593,0.0004084401,0.00007457879,0.0005076763,0.00027189765,0.0001780715],"domain_scores_gemma":[0.99021727,0.0071400786,0.00039934262,0.0010990992,0.00092152075,0.0002226557],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018973436,0.002280262,0.0010772962,0.0009851683,0.00063851767,0.0018763685,0.0031619964,0.0023130754,0.009732446],"category_scores_gemma":[0.015273831,0.0010571213,0.0013665616,0.00081231556,0.0013671161,0.004792671,0.0020196654,0.0051439917,0.003414466],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052825065,0.000856334,0.006370805,0.0006313867,0.00019252975,0.00021207215,0.00013447576,0.67055684,0.005577173,0.013865444,0.027438205,0.27363646],"study_design_scores_gemma":[0.00001950317,0.00003480093,0.00013271213,0.000012571711,0.000008306731,0.000009436352,0.000013657008,0.98544705,0.0013068966,0.012495995,0.00051453005,0.000004514319],"about_ca_topic_score_codex":0.0042974856,"about_ca_topic_score_gemma":0.014410008,"teacher_disagreement_score":0.009732446,"about_ca_system_score_codex":0.0020116053,"about_ca_system_score_gemma":0.002649927,"threshold_uncertainty_score":0.032558322},"labels":[],"label_agreement":null},{"id":"W2997847174","doi":"10.1609/aaai.v34i05.6430","title":"TreeGen: A Tree-Based Transformer Architecture for Code Generation","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":163,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"Computer science; Python (programming language); Code generation; Parsing; Programming language; Parse tree; Redundant code; Encoder; Source code; Abstract syntax tree; Artificial intelligence; Theoretical computer science; Operating system","score_opus":0.12118065944043592,"score_gpt":0.30576729455314144,"score_spread":0.1845866351127055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997847174","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014219388,0.00039002634,0.90267533,0.0002534108,0.000101648766,0.00020032741,0.00091935974,0.077011846,0.0042286096],"genre_scores_gemma":[0.22201927,0.0005611795,0.75683886,0.0005404636,0.00003942374,0.00041871908,0.005590188,0.004716355,0.009275631],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99963844,0.00006185934,0.000026164307,0.00012691853,0.00011206323,0.0000344991],"domain_scores_gemma":[0.9992675,0.00029264638,0.000048036447,0.00019821376,0.0001579552,0.000035680845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00061657804,0.00076530065,0.00035291,0.00083063764,0.00031468898,0.00065042556,0.0028368186,0.0008061489,0.006474998],"category_scores_gemma":[0.0024711534,0.0006115641,0.00076937285,0.000684381,0.00060817733,0.0027078448,0.001095465,0.0013896041,0.0022832693],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004714823,0.00031545968,0.0032305566,0.0006684947,0.00016045907,0.00034290308,0.00031101945,0.13527605,0.0318538,0.029253528,0.06003231,0.738084],"study_design_scores_gemma":[0.00006061793,0.00011610944,0.00038191333,0.000030211066,0.000053057025,0.00016483432,0.00002541597,0.9347256,0.02671577,0.01779464,0.019902812,0.00002895227],"about_ca_topic_score_codex":0.0072590723,"about_ca_topic_score_gemma":0.01185979,"teacher_disagreement_score":0.0072590723,"about_ca_system_score_codex":0.0010306264,"about_ca_system_score_gemma":0.0014593098,"threshold_uncertainty_score":0.021661043},"labels":[],"label_agreement":null},{"id":"W2997945246","doi":"10.5555/3370272.3370293","title":"On the Distribution of Test Smells in Open Source Android Applications: An Exploratory Study","year":2019,"lang":"en","type":"article","venue":"Espace ÉTS (ETS)","topic":"Software Engineering Research","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Code smell; Android (operating system); Computer science; Software quality; Unit testing; Empirical research; Software engineering; Software; Software development; Operating system","score_opus":0.020695644231862865,"score_gpt":0.28351608318632576,"score_spread":0.2628204389544629,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997945246","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9990005,0.00008073603,0.0004355754,0.00003161448,0.0000015362293,0.000020471029,0.00012599831,0.000014364238,0.0002891878],"genre_scores_gemma":[0.9987318,0.000100227626,0.00059708604,0.00002181808,0.000007737877,0.000033995104,0.0002688679,0.000013150879,0.00022532363],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9960795,0.001189439,0.00052082015,0.0006486914,0.0012245413,0.00033699177],"domain_scores_gemma":[0.92361975,0.04691074,0.018303547,0.0031072511,0.006380005,0.0016787638],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002912091,0.00032738197,0.00037323375,0.0045696152,0.00046837292,0.0010281791,0.0005125456,0.000730012,0.00091378833],"category_scores_gemma":[0.03077168,0.00026805952,0.0003641583,0.0025474085,0.0007845505,0.0014460937,0.0013924158,0.00066782406,0.00036090732],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024091676,0.00037603744,0.94185627,0.00020540191,0.00006605357,0.0014943857,0.018449638,0.00034406566,0.004165228,0.00017683125,0.00047653943,0.032148667],"study_design_scores_gemma":[0.000005806083,0.00035224945,0.9863892,0.00006704231,0.000025563093,0.0013951095,0.007897429,0.0014743712,0.0011866591,0.00013870929,0.0010376939,0.000030044286],"about_ca_topic_score_codex":0.0015235947,"about_ca_topic_score_gemma":0.0024589018,"teacher_disagreement_score":0.0045696152,"about_ca_system_score_codex":0.0004290534,"about_ca_system_score_gemma":0.00041363222,"threshold_uncertainty_score":0.015400767},"labels":[],"label_agreement":null},{"id":"W3000090264","doi":"10.1007/978-3-030-35510-4_1","title":"Together We Are Stronger: Evidence-Based Reflections on Industry-Academia Collaboration in Software Testing","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Queen's University; Queen's University Belfast","keywords":"Context (archaeology); Software testing; Software technical review; Empirical evidence; Software; Empirical research; Grey literature; Knowledge management; Engineering; Engineering management; Software engineering; Computer science; Engineering ethics; Software development; Political science; Software design; Geography","score_opus":0.07078751905955358,"score_gpt":0.3270566590593803,"score_spread":0.2562691399998267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3000090264","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017980948,0.009952505,0.0036741435,0.95466924,0.0011021405,0.000045187844,0.000063034095,0.000024681256,0.012488098],"genre_scores_gemma":[0.8426254,0.010458715,0.00927341,0.13306786,0.001522314,0.00026400673,0.00015410586,0.00013611445,0.0024980116],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.66150784,0.270097,0.015097968,0.010354158,0.033790514,0.009152518],"domain_scores_gemma":[0.07753843,0.8779472,0.010075264,0.008652809,0.019230703,0.0065555302],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.36031806,0.00094227074,0.0016319427,0.0059524374,0.01056326,0.033008885,0.0127638485,0.031250235,0.008613715],"category_scores_gemma":[0.58423096,0.0013020148,0.0010736268,0.008155404,0.04969482,0.048134137,0.02774896,0.04627666,0.0011741073],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069103955,0.0006757856,0.017249651,0.0038738162,0.0004878529,0.0023154898,0.23166579,0.0014006824,0.001091761,0.376665,0.108606584,0.2552766],"study_design_scores_gemma":[0.00037759225,0.0005565361,0.011929038,0.024077678,0.0003893757,0.0011423873,0.34398866,0.0014258659,0.0020953498,0.3907993,0.22290438,0.00031386825],"about_ca_topic_score_codex":0.009237882,"about_ca_topic_score_gemma":0.016551306,"teacher_disagreement_score":0.36031806,"about_ca_system_score_codex":0.016471952,"about_ca_system_score_gemma":0.038110733,"threshold_uncertainty_score":0.7888417},"labels":[],"label_agreement":null},{"id":"W3000135256","doi":"10.1109/ase.2019.00099","title":"CLCDSA: Cross Language Code Clone Detection using Syntactical Features and API Documentation","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":101,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Source code; Programming language; Compiler; Software; Software maintenance; clone (Java method); Artificial intelligence; Natural language processing; Software development","score_opus":0.011631929083889644,"score_gpt":0.3243410991509543,"score_spread":0.31270917006706467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3000135256","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32861543,0.0013465749,0.58469826,0.00044986318,0.00013211381,0.00064461137,0.0033702212,0.07653163,0.0042113406],"genre_scores_gemma":[0.7003262,0.000299186,0.28442287,0.00032969174,0.00003967899,0.00035241604,0.0073929047,0.00096996006,0.005867134],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9978036,0.0002249552,0.00019268358,0.00066732225,0.00097044045,0.00014087127],"domain_scores_gemma":[0.9914927,0.0024259414,0.0020427285,0.0013387945,0.0024483444,0.00025154147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014756803,0.0012639199,0.0009530717,0.0057478617,0.0005605435,0.0015960162,0.0019502168,0.0013755255,0.0010229328],"category_scores_gemma":[0.008693219,0.0005534115,0.0014737543,0.0021947548,0.00071884826,0.0023053158,0.0019976716,0.0012628849,0.0010248227],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057315454,0.0006346379,0.1301445,0.000718561,0.00042971203,0.0009767277,0.00081533584,0.028891476,0.047318097,0.003570007,0.0145326415,0.7713952],"study_design_scores_gemma":[0.00007737099,0.00043762018,0.041926134,0.00009256572,0.00017491815,0.0012523754,0.00021261195,0.8703853,0.06506837,0.004923345,0.015325734,0.00012365264],"about_ca_topic_score_codex":0.012207121,"about_ca_topic_score_gemma":0.013472166,"teacher_disagreement_score":0.012207121,"about_ca_system_score_codex":0.0011132238,"about_ca_system_score_gemma":0.0020738433,"threshold_uncertainty_score":0.024272144},"labels":[],"label_agreement":null},{"id":"W3000602293","doi":"10.1109/tcad.2020.2966448","title":"Searching for Bugs Using Probabilistic Suspect Implications","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Computer-Aided Design of Integrated Circuits and Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Debugging; Suspect; Computer science; Probabilistic logic; Set (abstract data type); Machine learning; Hyperparameter; Algorithmic program debugging; Artificial intelligence; Software bug; Data mining; Programming language; Software","score_opus":0.10228609820325334,"score_gpt":0.28921066827009445,"score_spread":0.1869245700668411,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3000602293","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26458532,0.00072647585,0.7234868,0.001298296,0.00004053731,0.00015811205,0.0008179267,0.0061023124,0.0027842321],"genre_scores_gemma":[0.8456114,0.00015577633,0.15227777,0.000121609024,0.000019164225,0.00005661129,0.0006508368,0.0001285592,0.0009781957],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99899143,0.00023087513,0.00006741744,0.0002607811,0.00036944263,0.00008014683],"domain_scores_gemma":[0.990905,0.006629077,0.00092642,0.000520407,0.00088612904,0.00013292246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001419399,0.00091616437,0.00061450386,0.0034037298,0.0005475837,0.001138258,0.0014264095,0.0009936255,0.0026811296],"category_scores_gemma":[0.015365427,0.0005506764,0.0008086606,0.0011162934,0.0006865096,0.0025874896,0.0011371224,0.0010669698,0.00039574245],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066514505,0.00029403198,0.09512095,0.00041662634,0.00012476233,0.00083407795,0.0007515766,0.36745003,0.009053805,0.015020754,0.0056636785,0.5046045],"study_design_scores_gemma":[0.000022653727,0.00006948395,0.0023124237,0.0000363985,0.00003130648,0.00018743446,0.00007343568,0.9786734,0.0025658032,0.015308288,0.0007043138,0.000015019436],"about_ca_topic_score_codex":0.0048921513,"about_ca_topic_score_gemma":0.01153109,"teacher_disagreement_score":0.0048921513,"about_ca_system_score_codex":0.00082289474,"about_ca_system_score_gemma":0.0014141666,"threshold_uncertainty_score":0.009727359},"labels":[],"label_agreement":null},{"id":"W3001461005","doi":"10.1109/tse.2020.2967380","title":"A Machine Learning Approach to Improve the Detection of CI Skip Commits","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Commit; Computer science; Machine learning; Facilitator; Decision tree; Artificial intelligence; Tree (set theory); Software; Popularity; Data mining; Software engineering; Database; Operating system","score_opus":0.01601098347251498,"score_gpt":0.21960679703534175,"score_spread":0.20359581356282677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3001461005","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43202573,0.0020081345,0.5425647,0.0017149051,0.00046226443,0.0006231099,0.003948612,0.011754049,0.00489842],"genre_scores_gemma":[0.8252323,0.00021736146,0.16818781,0.00024112668,0.00014276954,0.0002516361,0.0038253646,0.000074876225,0.001826753],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99659854,0.00053967495,0.00055422547,0.0010148507,0.00085871154,0.0004340363],"domain_scores_gemma":[0.98483014,0.0071334722,0.0023233213,0.00094643945,0.004294467,0.00047207097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035416132,0.0014408735,0.0012936488,0.008532834,0.00080309424,0.0015024858,0.001996175,0.00153225,0.0010546292],"category_scores_gemma":[0.014653332,0.00030022053,0.0009999755,0.0041996865,0.0004294133,0.0015758182,0.0008356221,0.002014086,0.0011155652],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043368968,0.0016132161,0.18168782,0.00033279645,0.00029681987,0.0007884841,0.00031985287,0.062693655,0.010493075,0.0014095287,0.012703349,0.7272276],"study_design_scores_gemma":[0.00002001765,0.00019658078,0.018183578,0.000044694876,0.00006850734,0.000312341,0.00014153782,0.9705283,0.0068273204,0.0017111327,0.0019271838,0.000038689708],"about_ca_topic_score_codex":0.00884106,"about_ca_topic_score_gemma":0.008269885,"teacher_disagreement_score":0.00884106,"about_ca_system_score_codex":0.00094023946,"about_ca_system_score_gemma":0.0017709008,"threshold_uncertainty_score":0.018730104},"labels":[],"label_agreement":null},{"id":"W3001783472","doi":"10.1007/s10664-020-09807-w","title":"Ammonia: an approach for deriving project-specific bug patterns","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Java; Software bug; Software regression; Security bug; Software development; Source code; Limiting; Pointer (user interface); Software maintenance","score_opus":0.0743664019156529,"score_gpt":0.30118913813212417,"score_spread":0.22682273621647125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3001783472","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037677683,0.00011981761,0.9351523,0.00027787662,0.00006537372,0.0007223539,0.002360571,0.021061191,0.0025628302],"genre_scores_gemma":[0.099550895,0.00010875428,0.89341533,0.000091975024,0.000027611064,0.00067192764,0.0032858134,0.0010729098,0.0017747233],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99536514,0.0009905199,0.0006477621,0.0011302839,0.0016611929,0.0002050526],"domain_scores_gemma":[0.98513675,0.0055111265,0.0027858615,0.0027986409,0.0033766003,0.0003911018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038281824,0.0014707645,0.00087465433,0.008736753,0.0009987787,0.0023339852,0.0017978718,0.0012194605,0.002756398],"category_scores_gemma":[0.025865888,0.0010206068,0.0014094635,0.00497124,0.00066346175,0.0024532252,0.0024485502,0.0012561668,0.0014528795],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003287719,0.00049540255,0.10563283,0.0013833855,0.00037945432,0.0014754318,0.004136307,0.021953043,0.033638775,0.014658751,0.016673513,0.7992443],"study_design_scores_gemma":[0.00017010723,0.0006991554,0.049265824,0.00067390443,0.0005906918,0.0033378527,0.0022599152,0.793428,0.045848876,0.03036886,0.073033385,0.0003234141],"about_ca_topic_score_codex":0.0044375323,"about_ca_topic_score_gemma":0.006944282,"teacher_disagreement_score":0.008736753,"about_ca_system_score_codex":0.00073885446,"about_ca_system_score_gemma":0.0031267728,"threshold_uncertainty_score":0.020245612},"labels":[],"label_agreement":null},{"id":"W3003493628","doi":"10.4018/978-1-7998-1863-2.ch007","title":"Data in DevOps and Its Importance in Code Analytics","year":2020,"lang":"en","type":"book-chapter","venue":"Advances in systems analysis, software engineering, and high performance computing book series","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Leukemia & Lymphoma Society of Canada; Cisco Systems (Canada)","funders":"","keywords":"Commit; DevOps; Computer science; Software; Software engineering; Data science; Computer security; Database; Operating system","score_opus":0.01722413630931374,"score_gpt":0.24818583137596645,"score_spread":0.2309616950666527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3003493628","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013262737,0.16178747,0.21132603,0.07005871,0.008020842,0.00021099801,0.013342432,0.0037987053,0.5181921],"genre_scores_gemma":[0.12524068,0.24972633,0.28915283,0.014923212,0.01104549,0.00056756247,0.026275462,0.008349608,0.27471873],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.997951,0.00041730257,0.00013627991,0.00033341057,0.001097211,0.00006476565],"domain_scores_gemma":[0.9820011,0.012993088,0.0005203258,0.0019268725,0.0020778505,0.00048065532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028268946,0.0007260888,0.0006502117,0.007573735,0.0014188612,0.009988194,0.0011837932,0.0014327433,0.02287342],"category_scores_gemma":[0.016995903,0.00086978957,0.0007088069,0.016346598,0.0034212375,0.017946297,0.003666538,0.005311681,0.007707428],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000339239,0.000029571858,0.0016095132,0.00081613194,0.000016885828,0.00012890811,0.0016590391,0.0014739948,0.00066554395,0.5033069,0.16437933,0.3258802],"study_design_scores_gemma":[0.0000026608295,0.0000131167735,0.0014881777,0.0011334671,0.0000058297937,0.00032119663,0.00056140503,0.001774873,0.00054521393,0.22203545,0.7720871,0.000031452753],"about_ca_topic_score_codex":0.0032335832,"about_ca_topic_score_gemma":0.0036488064,"teacher_disagreement_score":0.02287342,"about_ca_system_score_codex":0.0023453827,"about_ca_system_score_gemma":0.0017026553,"threshold_uncertainty_score":0.07651925},"labels":[],"label_agreement":null},{"id":"W3004067286","doi":"","title":"Detecting and Correcting Typing Errors in DBpedia.","year":2019,"lang":"en","type":"article","venue":"Knowledge Discovery and Data Mining","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Natural language processing; Typing; Information retrieval; Artificial intelligence; Speech recognition","score_opus":0.03783106692397872,"score_gpt":0.3066728428562663,"score_spread":0.26884177593228753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3004067286","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17467025,0.0071981153,0.59210545,0.006023504,0.0058962065,0.0015230237,0.09733302,0.09329607,0.021954427],"genre_scores_gemma":[0.24770828,0.0023263479,0.61867094,0.0022107353,0.000493262,0.00042621876,0.11094508,0.008071883,0.009147318],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9836586,0.0044754134,0.0023831974,0.0033037493,0.0053934716,0.0007855901],"domain_scores_gemma":[0.9362258,0.029126272,0.0044845985,0.014039313,0.015169177,0.00095496455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0095052,0.0016865684,0.001533593,0.008312087,0.0026687917,0.00629591,0.0031526107,0.0026400993,0.0022337746],"category_scores_gemma":[0.06935366,0.0011655254,0.0014235914,0.0067532524,0.001005926,0.0064822594,0.0049691177,0.0030078043,0.004261702],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015028688,0.0012854362,0.06061647,0.00450612,0.0011086265,0.004624105,0.005694993,0.008762404,0.030614763,0.0165613,0.29534572,0.5693772],"study_design_scores_gemma":[0.0003020753,0.00041544708,0.029871752,0.0036178066,0.0014887219,0.005351347,0.0071592345,0.14825702,0.18153411,0.07943864,0.54189223,0.00067164155],"about_ca_topic_score_codex":0.008348448,"about_ca_topic_score_gemma":0.012223271,"teacher_disagreement_score":0.0095052,"about_ca_system_score_codex":0.0009587288,"about_ca_system_score_gemma":0.0051576677,"threshold_uncertainty_score":0.05026889},"labels":[],"label_agreement":null},{"id":"W3004570974","doi":"10.1007/s10664-019-09781-y","title":"How bugs are born: a model to identify how bugs are introduced in software components","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":74,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Waterloo","funders":"H2020 Industrial Leadership; Ministerio de Asuntos Económicos y Transformación Digital, Gobierno de España; Nederlandse Organisatie voor Wetenschappelijk Onderzoek","keywords":"Software bug; Software regression; Computer science; False positive paradox; Source lines of code; Software; Source code; Open source; Debugging; Software maintenance; Snapshot (computer storage); Code (set theory); Security bug; Data mining; Software development; Software quality; Programming language; Machine learning; Database; Operating system; Set (abstract data type); Software security assurance","score_opus":0.062116288198003355,"score_gpt":0.2983208692382943,"score_spread":0.23620458104029096,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3004570974","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.62440807,0.0021242509,0.35925683,0.0036663774,0.00010529376,0.00035073224,0.0028883952,0.003922355,0.003277577],"genre_scores_gemma":[0.9424195,0.00027681788,0.05380355,0.0002613464,0.000043990505,0.00011864904,0.0021675325,0.0001172587,0.00079130486],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958449,0.0012613453,0.00038031166,0.001313354,0.0008557821,0.00034424895],"domain_scores_gemma":[0.93993104,0.04215767,0.008307773,0.0034703333,0.004912148,0.001221157],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070214965,0.0013356216,0.0010450677,0.00929131,0.0009392814,0.0038458325,0.0021303005,0.0030915784,0.001811786],"category_scores_gemma":[0.049292944,0.0007353934,0.0018044963,0.003447786,0.0025210644,0.006107893,0.0023212172,0.0016079808,0.0007701086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013689456,0.0005830828,0.5505127,0.001018617,0.00052371534,0.0011827888,0.0034109924,0.26920176,0.006679441,0.023908623,0.012215618,0.12939379],"study_design_scores_gemma":[0.000047697282,0.00014182208,0.033474814,0.0000883963,0.00007424544,0.00042687793,0.00045079426,0.9415368,0.0012370471,0.020761326,0.0017055571,0.000054610784],"about_ca_topic_score_codex":0.011138844,"about_ca_topic_score_gemma":0.009765143,"teacher_disagreement_score":0.011138844,"about_ca_system_score_codex":0.0023772153,"about_ca_system_score_gemma":0.0015975679,"threshold_uncertainty_score":0.037133694},"labels":[],"label_agreement":null},{"id":"W3004796001","doi":"10.1016/j.infsof.2020.106277","title":"Mining API usage scenarios from stack overflow","year":2020,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Polytechnique Montréal","funders":"","keywords":"Documentation; Computer science; Application programming interface; Task (project management); Software documentation; Source code; Code (set theory); Software; Code review; Programming language; Internal documentation; Software engineering; World Wide Web; Database; Software development; Software quality; Software development process; Engineering; Software construction; Set (abstract data type)","score_opus":0.014472455903386122,"score_gpt":0.23106877270585893,"score_spread":0.21659631680247282,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3004796001","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9750123,0.00054948224,0.008431274,0.00020853945,0.000036020137,0.00010149608,0.010687754,0.0027057054,0.0022674673],"genre_scores_gemma":[0.9599191,0.00035076228,0.01828081,0.00006927295,0.000032430107,0.000084081046,0.020292992,0.00026914352,0.0007015851],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99799347,0.0002284499,0.00025257887,0.00039213995,0.0009227443,0.00021062137],"domain_scores_gemma":[0.99118257,0.0045088325,0.0015237766,0.0009148306,0.0013031537,0.0005667812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008886356,0.0011273716,0.00046633216,0.008561466,0.0007333873,0.0014171434,0.0010057492,0.0013147666,0.0009656119],"category_scores_gemma":[0.011779192,0.0004736879,0.0013027441,0.0056416113,0.00034353294,0.003096105,0.0010420717,0.0009170902,0.00052749855],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017006417,0.0008545224,0.6812137,0.0015558656,0.00083433505,0.010802006,0.0024663664,0.02266306,0.037289996,0.0050063543,0.020818328,0.21479495],"study_design_scores_gemma":[0.00010311227,0.0007007593,0.37606058,0.0005039206,0.0009781742,0.012979804,0.0038999075,0.5211865,0.04190929,0.012044973,0.02936497,0.0002679654],"about_ca_topic_score_codex":0.003867164,"about_ca_topic_score_gemma":0.0065974584,"teacher_disagreement_score":0.008561466,"about_ca_system_score_codex":0.0004810702,"about_ca_system_score_gemma":0.0009496666,"threshold_uncertainty_score":0.007689297},"labels":[],"label_agreement":null},{"id":"W3005009401","doi":"","title":"A Toolchain to Produce Correct-by-Construction OCaml Programs","year":2018,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure","funders":"","keywords":"Toolchain; Programming language; Computer science; Operating system; Software","score_opus":0.01367944044203232,"score_gpt":0.23664346775261194,"score_spread":0.22296402731057963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3005009401","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0135983145,0.00006681851,0.9739577,0.0002819363,0.00010873729,0.00011701085,0.0001114807,0.007914536,0.0038434733],"genre_scores_gemma":[0.23443954,0.00023775612,0.75236225,0.0002768415,0.000056129313,0.00041456407,0.000627458,0.0049944064,0.006591067],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9926077,0.0019168726,0.00066768925,0.000987306,0.0029711707,0.0008491287],"domain_scores_gemma":[0.9829202,0.0074682683,0.0006719375,0.0056897034,0.0028512168,0.0003986099],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00661246,0.00095030887,0.0007333415,0.0012787166,0.000997331,0.0020983466,0.0018162592,0.001459797,0.008253702],"category_scores_gemma":[0.02897818,0.0010982733,0.0016128533,0.000759491,0.0025962007,0.0038713398,0.0051816325,0.0025334142,0.0034292762],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051082205,0.00047606928,0.0043875845,0.0012554969,0.00018605459,0.0018993111,0.002295649,0.058385663,0.08753721,0.5493579,0.012179408,0.2815288],"study_design_scores_gemma":[0.00033765635,0.0004712741,0.001029615,0.00049787434,0.00024320172,0.0012349956,0.00041085857,0.28368664,0.2547044,0.354467,0.102705866,0.0002105593],"about_ca_topic_score_codex":0.0010770787,"about_ca_topic_score_gemma":0.0011104366,"teacher_disagreement_score":0.008253702,"about_ca_system_score_codex":0.0007650265,"about_ca_system_score_gemma":0.002960915,"threshold_uncertainty_score":0.034970462},"labels":[],"label_agreement":null},{"id":"W3006039974","doi":"10.1109/tse.2020.2973997","title":"ConfigMiner: Identifying the Appropriate Configuration Options for Config-Related User Questions by Mining Online Forums","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Task (project management); Rank (graph theory); Software; State (computer science); World Wide Web; Information retrieval; Data science; Programming language; Engineering; Systems engineering","score_opus":0.02648402744795977,"score_gpt":0.26287446803427095,"score_spread":0.23639044058631117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3006039974","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50239027,0.004200014,0.29989877,0.001851212,0.0002906933,0.0017904658,0.03412071,0.14622231,0.009235489],"genre_scores_gemma":[0.7116198,0.0006108514,0.2442765,0.00072008977,0.00014212083,0.0011914824,0.034924537,0.0021073166,0.00440734],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99569076,0.0016662066,0.0003723279,0.001051867,0.0010210895,0.00019768858],"domain_scores_gemma":[0.9784386,0.016293138,0.0018568527,0.001661145,0.0012157677,0.0005345025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003935153,0.0027034509,0.0010993655,0.0047783325,0.00053533056,0.0015861146,0.0016987094,0.0017226755,0.0036160615],"category_scores_gemma":[0.021621117,0.0005563741,0.0009364822,0.0012418655,0.00046993926,0.004137794,0.0016362866,0.0011635771,0.003108296],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021979301,0.0018509811,0.122926906,0.0033517177,0.00032708078,0.002035319,0.0042471145,0.011152651,0.044014018,0.0026573292,0.08544515,0.71979374],"study_design_scores_gemma":[0.00044214845,0.001994992,0.10418809,0.00079532654,0.00032387523,0.004908059,0.0043585333,0.7026992,0.06505839,0.01712997,0.09764477,0.00045669934],"about_ca_topic_score_codex":0.0008316546,"about_ca_topic_score_gemma":0.0024849747,"teacher_disagreement_score":0.0047783325,"about_ca_system_score_codex":0.0005138492,"about_ca_system_score_gemma":0.00062569516,"threshold_uncertainty_score":0.020811379},"labels":[],"label_agreement":null},{"id":"W3006884129","doi":"10.1145/3373087.3375324","title":"Programming Abstractions for Configurable Hardware","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Toolchain; Programming language; Domain-specific language; Compiler; Programming paradigm; Computer architecture; Programmer; Software","score_opus":0.0482413803972899,"score_gpt":0.2919442547878401,"score_spread":0.2437028743905502,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3006884129","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004524446,0.0005073597,0.98488533,0.00036106052,0.000045082943,0.000059839622,0.000099020624,0.0012985349,0.00821947],"genre_scores_gemma":[0.10113513,0.0019157258,0.8865825,0.00038284782,0.00007051537,0.0003846359,0.0003783577,0.00071922527,0.008431018],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987795,0.00032549904,0.00013223627,0.00017622921,0.00047869314,0.00010788776],"domain_scores_gemma":[0.99813604,0.0007495813,0.00015743967,0.0006992052,0.00019287963,0.00006482172],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018983353,0.0005979095,0.00028569085,0.00072125794,0.00046856632,0.0024932234,0.0013886159,0.00070340076,0.005688078],"category_scores_gemma":[0.0036299885,0.00067148,0.000899383,0.00074420945,0.0016446335,0.0033710953,0.0018449796,0.0029651185,0.0014572283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028958062,0.00002704055,0.00040990717,0.00040520288,0.0000249601,0.0000818713,0.00048190582,0.02026756,0.0069173123,0.9101478,0.0033385307,0.057869032],"study_design_scores_gemma":[0.000048726568,0.000086922344,0.00040832238,0.00047083796,0.00007311279,0.00039741225,0.00024389903,0.1477081,0.01532564,0.52488,0.31031108,0.00004605119],"about_ca_topic_score_codex":0.0006465643,"about_ca_topic_score_gemma":0.0012697405,"teacher_disagreement_score":0.005688078,"about_ca_system_score_codex":0.0010069981,"about_ca_system_score_gemma":0.0016182639,"threshold_uncertainty_score":0.019028485},"labels":[],"label_agreement":null},{"id":"W3007452449","doi":"10.1002/smr.2255","title":"Bad smell detection using quality metrics and refactoring opportunities","year":2020,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code smell; Code refactoring; Computer science; Maintainability; False positive paradox; Process (computing); Quality (philosophy); Software quality; Set (abstract data type); Technical debt; Software maintenance; Software; Code (set theory); Source code; Software engineering; Software system; Artificial intelligence; Software development; Programming language","score_opus":0.1255498175057313,"score_gpt":0.3338147225408901,"score_spread":0.20826490503515882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3007452449","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56964105,0.0007056942,0.4220247,0.00035743066,0.0000615006,0.00032239675,0.0004253513,0.0042035673,0.0022583138],"genre_scores_gemma":[0.80039936,0.000099805264,0.19837742,0.000029350113,0.0000139277,0.00007662018,0.0004199233,0.00013516213,0.0004484616],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9884609,0.0029128422,0.0012494486,0.0012413887,0.00564541,0.00049007847],"domain_scores_gemma":[0.924802,0.025372433,0.020225234,0.005068013,0.022801025,0.001731335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00697754,0.0014041394,0.0010367552,0.011536205,0.00064859824,0.00219256,0.0010677737,0.0009624009,0.00058042683],"category_scores_gemma":[0.03600485,0.00047874427,0.0009318135,0.003288652,0.0006135795,0.002094208,0.0016985631,0.00089284225,0.00026535676],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005855277,0.00077095826,0.3939757,0.0009195363,0.00040193138,0.0008669283,0.0018224365,0.034715183,0.070658,0.003090045,0.0021793307,0.4900143],"study_design_scores_gemma":[0.00009388389,0.0012619058,0.16379125,0.00029047354,0.0003097461,0.0012619388,0.0008318957,0.72351295,0.099117994,0.005176683,0.0040496504,0.00030161405],"about_ca_topic_score_codex":0.003773572,"about_ca_topic_score_gemma":0.0049892766,"teacher_disagreement_score":0.011536205,"about_ca_system_score_codex":0.0009975315,"about_ca_system_score_gemma":0.0013565562,"threshold_uncertainty_score":0.036901236},"labels":[],"label_agreement":null},{"id":"W3008252796","doi":"10.1109/icmla.2019.00096","title":"Feature Changes in Source Code for Commit Classification Into Maintenance Activities","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Commit; Software maintenance; Computer science; Boosting (machine learning); Backporting; Metric (unit); Software development; Software engineering; Feature (linguistics); Software; Baseline (sea); Software bug; Machine learning; Data mining; Artificial intelligence; Software construction; Database; Engineering; Programming language","score_opus":0.022525687730514006,"score_gpt":0.2739560130404134,"score_spread":0.2514303253098994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3008252796","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8382034,0.002292917,0.121224456,0.00064925686,0.00042746385,0.00054768776,0.012052441,0.018230502,0.0063718935],"genre_scores_gemma":[0.8992969,0.00032225635,0.07647639,0.00007941429,0.00010624992,0.00019548093,0.018665789,0.00034697985,0.004510485],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99883574,0.0000918834,0.00008809969,0.00037008198,0.00046253213,0.00015155284],"domain_scores_gemma":[0.9964407,0.00083500985,0.0006814554,0.0006020905,0.0011348967,0.0003058389],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00070245087,0.00084770756,0.0005686529,0.0050768806,0.00047172714,0.00088304887,0.0008647681,0.00083300605,0.002145537],"category_scores_gemma":[0.0051449044,0.00017234497,0.0007786123,0.0023105324,0.00025478212,0.0011832054,0.00073275156,0.0010167195,0.002097104],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006542101,0.0007176806,0.20574889,0.00039299054,0.00012779035,0.00043744262,0.00032401102,0.0071854657,0.021038624,0.0007142332,0.025148245,0.7375105],"study_design_scores_gemma":[0.000103680715,0.0007905645,0.374444,0.00020855146,0.00027596057,0.0017081667,0.0006483608,0.5425179,0.044735532,0.0030843772,0.03136398,0.00011889954],"about_ca_topic_score_codex":0.0041917074,"about_ca_topic_score_gemma":0.0100085465,"teacher_disagreement_score":0.0050768806,"about_ca_system_score_codex":0.00045202405,"about_ca_system_score_gemma":0.0008069463,"threshold_uncertainty_score":0.008334577},"labels":[],"label_agreement":null},{"id":"W3008292200","doi":"10.1109/icmla.2019.00195","title":"Software Fault Prediction Based on Fault Probability and Impact","year":2019,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Computer science; Machine learning; Naive Bayes classifier; Fault (geology); Random forest; Artificial intelligence; Data mining; Software; Support vector machine; Software system; Perceptron; Artificial neural network","score_opus":0.013471463074528676,"score_gpt":0.2608939757624029,"score_spread":0.24742251268787424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3008292200","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8139529,0.0005473043,0.18037184,0.00017656942,0.000038537684,0.00009164635,0.0012937421,0.0017605725,0.0017668388],"genre_scores_gemma":[0.98627496,0.00008220587,0.012596188,0.0000075724843,0.000009532397,0.000020156407,0.00065161585,0.00003653754,0.00032125346],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99867755,0.00022434183,0.00012936819,0.00030548192,0.00053975784,0.00012343291],"domain_scores_gemma":[0.9876627,0.007983463,0.0015931448,0.00067758025,0.0018377108,0.00024547204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015479457,0.000935462,0.00071714114,0.004239642,0.0002314632,0.0007483838,0.00063574663,0.0006998282,0.0012275582],"category_scores_gemma":[0.01346801,0.00026363565,0.00087981956,0.0014960178,0.00034486124,0.0013498302,0.0005114682,0.0006261619,0.00050866825],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003838857,0.00017790901,0.17544563,0.00012374757,0.00013465338,0.00022195764,0.000098384495,0.7332872,0.003186758,0.0007240926,0.00080998405,0.085405715],"study_design_scores_gemma":[0.000005029485,0.00006999387,0.02899407,0.000010608286,0.000023093107,0.000085930355,0.000014796273,0.967767,0.0016153554,0.0012362713,0.00016484968,0.000012960052],"about_ca_topic_score_codex":0.0051531047,"about_ca_topic_score_gemma":0.004238042,"teacher_disagreement_score":0.0051531047,"about_ca_system_score_codex":0.000687901,"about_ca_system_score_gemma":0.00039571113,"threshold_uncertainty_score":0.010246277},"labels":[],"label_agreement":null},{"id":"W3009319801","doi":"10.1145/3328778.3372601","title":"Analyzing CS1 Student Code Using Code Embeddings","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Code (set theory); Programming language; Source code; Theoretical computer science; Parallel computing","score_opus":0.055093297533597416,"score_gpt":0.33990359395726133,"score_spread":0.2848102964236639,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3009319801","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6737003,0.00036598396,0.31650212,0.0008431631,0.00014639909,0.00012643135,0.0022968606,0.0033428115,0.0026759785],"genre_scores_gemma":[0.91022044,0.00012129414,0.08272339,0.00007512709,0.000042453205,0.00013052623,0.0036641024,0.00029854995,0.00272408],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982091,0.0005477099,0.00012291933,0.0004034473,0.0005797443,0.0001370804],"domain_scores_gemma":[0.9888842,0.005063366,0.0015373618,0.0011675016,0.0028732289,0.00047436293],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014886822,0.0008173301,0.00035521944,0.0021860152,0.00039774814,0.0012514979,0.00068961416,0.0007996925,0.0017965304],"category_scores_gemma":[0.01864592,0.00021463177,0.000514138,0.0019307634,0.000622229,0.0019691244,0.0012386121,0.0013912388,0.0007520976],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096648117,0.0010326825,0.14071593,0.0005584228,0.00016946145,0.00064163405,0.002185261,0.19192703,0.022113413,0.018174762,0.018179199,0.6033358],"study_design_scores_gemma":[0.000027508928,0.00028409736,0.013416925,0.00004936185,0.000019398281,0.00019589739,0.000570571,0.9429923,0.012182933,0.02528985,0.004920921,0.00005014124],"about_ca_topic_score_codex":0.001969281,"about_ca_topic_score_gemma":0.0031298997,"teacher_disagreement_score":0.0021860152,"about_ca_system_score_codex":0.00079493,"about_ca_system_score_gemma":0.0008648435,"threshold_uncertainty_score":0.007872999},"labels":[],"label_agreement":null},{"id":"W3009853975","doi":"10.1002/smr.2250","title":"Guidelines for evaluating bug‐assignment research","year":2020,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Faculty of Graduate Studies and Research, University of Alberta; Alberta Innovates - Technology Futures","keywords":"Computer science; Metric (unit); Ranking (information retrieval); Empirical research; Task (project management); Set (abstract data type); Data science; Data mining; Information retrieval; Statistics; Systems engineering; Mathematics; Engineering; Operations management","score_opus":0.28841862183190226,"score_gpt":0.4726799934303204,"score_spread":0.18426137159841816,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3009853975","genre_codex":"methods","genre_gemma":"methods","domain_codex":"methods","domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047015693,0.045569476,0.7464631,0.011733907,0.0025048296,0.092631035,0.010509613,0.006219345,0.037352957],"genre_scores_gemma":[0.055719562,0.0033176665,0.8557391,0.0012233597,0.00023835791,0.08007022,0.002315805,0.00045792462,0.0009179771],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.34285417,0.46980295,0.11128271,0.009128693,0.06487314,0.0020584257],"domain_scores_gemma":[0.11482825,0.6000945,0.059653107,0.038753048,0.18372384,0.0029473118],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.46915984,0.004148416,0.0055996105,0.04623996,0.0043127625,0.01351479,0.009125081,0.005672105,0.0056046606],"category_scores_gemma":[0.75546783,0.003114385,0.0061329748,0.027859954,0.0061240843,0.010115265,0.007026958,0.005002617,0.0043403376],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036096217,0.0023293337,0.03469468,0.060443435,0.0027382337,0.00065244816,0.017348152,0.007847033,0.011433183,0.048321303,0.08179494,0.72878766],"study_design_scores_gemma":[0.0067636063,0.009742638,0.09105337,0.15150332,0.00680815,0.0011658211,0.022174653,0.058832847,0.05737302,0.15811582,0.43425933,0.002207523],"about_ca_topic_score_codex":0.004802534,"about_ca_topic_score_gemma":0.007972544,"teacher_disagreement_score":0.53084016,"about_ca_system_score_codex":0.008857794,"about_ca_system_score_gemma":0.017137494,"threshold_uncertainty_score":0.65462047},"labels":[],"label_agreement":null},{"id":"W3010658448","doi":"10.1002/smr.2260","title":"Fuzzy case‐based‐reasoning‐based imputation for incomplete data in software engineering repositories","year":2020,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Missing data; Imputation (statistics); Categorical variable; Data mining; Computer science; Software; Fuzzy logic; Reuse; Machine learning; Artificial intelligence; Engineering","score_opus":0.03250375646217073,"score_gpt":0.29706369706087565,"score_spread":0.2645599405987049,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3010658448","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07303965,0.00047918278,0.9225016,0.0006858653,0.00007487127,0.00020378982,0.00042297773,0.0007852689,0.0018067722],"genre_scores_gemma":[0.61625195,0.00032404132,0.3811141,0.000144615,0.00004163919,0.0002524503,0.0010408879,0.00005975436,0.0007705331],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9827855,0.008271579,0.0016507044,0.0014775523,0.0052871024,0.00052755204],"domain_scores_gemma":[0.9356114,0.03999206,0.006050608,0.009574772,0.008197777,0.0005732951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021758065,0.00074997405,0.0018071266,0.0065005207,0.0015807374,0.0033581997,0.003949175,0.0019386791,0.0021712407],"category_scores_gemma":[0.07590329,0.00070130633,0.0022192313,0.006335717,0.0013682741,0.0043665324,0.0024249973,0.001865039,0.0005492689],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080995366,0.00077666284,0.041404657,0.00089434447,0.0008220514,0.0016932762,0.0019163506,0.516613,0.002963771,0.034623206,0.006317908,0.39116487],"study_design_scores_gemma":[0.000043258577,0.00006671569,0.003768181,0.00017266325,0.00010246108,0.00027901735,0.0002815406,0.96517664,0.0032659161,0.024745097,0.002036199,0.00006229805],"about_ca_topic_score_codex":0.009928925,"about_ca_topic_score_gemma":0.0070818653,"teacher_disagreement_score":0.021758065,"about_ca_system_score_codex":0.0022551129,"about_ca_system_score_gemma":0.0030124278,"threshold_uncertainty_score":0.11506903},"labels":[],"label_agreement":null},{"id":"W3010661194","doi":"10.1109/tifs.2020.2980190","title":"<i>CPA</i>: Accurate <i>C</i>ross-<i>P</i>latform Binary <i>A</i>uthorship Characterization Using LDA","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Information Forensics and Security","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"United Arab Emirates University","keywords":"Computer science; Binary number; Code refactoring; Binary code; Robustness (evolution); Compiler; Code (set theory); Coding (social sciences); Artificial intelligence; Natural language processing; Information retrieval; Set (abstract data type); Programming language; Software; Arithmetic","score_opus":0.028526798185236926,"score_gpt":0.24928584961411068,"score_spread":0.22075905142887375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3010661194","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04023239,0.0009140787,0.64207155,0.0016292061,0.0011945637,0.00068357895,0.048063103,0.23792703,0.027284529],"genre_scores_gemma":[0.3078514,0.0005873502,0.5543268,0.0007856212,0.0005191537,0.00095802353,0.08911522,0.01250446,0.033351954],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9971885,0.0003343083,0.00026815204,0.0009327357,0.0010485899,0.00022766039],"domain_scores_gemma":[0.9895064,0.001846532,0.001001803,0.0037931309,0.0035031477,0.00034901543],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020762843,0.0020136167,0.0011265079,0.006799058,0.0016302817,0.004685813,0.0016240275,0.001800301,0.023185372],"category_scores_gemma":[0.01739396,0.00077974267,0.00132302,0.003265583,0.0006934402,0.0039620916,0.0031034462,0.0022078382,0.03754347],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004560386,0.00014198889,0.019578036,0.00041504277,0.00013086807,0.00027623563,0.00023332602,0.006866023,0.015093515,0.0073806145,0.36202568,0.5874026],"study_design_scores_gemma":[0.000071341725,0.00008934801,0.019171482,0.00017944598,0.00006312382,0.0009410254,0.00039136637,0.6685562,0.071611,0.028832547,0.2099006,0.0001926087],"about_ca_topic_score_codex":0.006615321,"about_ca_topic_score_gemma":0.0114467405,"teacher_disagreement_score":0.023185372,"about_ca_system_score_codex":0.0012805875,"about_ca_system_score_gemma":0.0017954594,"threshold_uncertainty_score":0.07756281},"labels":[],"label_agreement":null},{"id":"W3012478566","doi":"10.1145/1035292.1029002","title":"Recovering binary class relationships","year":2004,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Programming language; Class (philosophy); Binary number; Theoretical computer science; Software; Object-oriented programming; Programming complexity; Software development; Software construction; Artificial intelligence","score_opus":0.04904132694768226,"score_gpt":0.2737487482202274,"score_spread":0.22470742127254512,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3012478566","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.087930776,0.0007339426,0.8930515,0.0011024462,0.000207749,0.0001228214,0.0012042788,0.0036281706,0.012018352],"genre_scores_gemma":[0.43575108,0.0004433889,0.5501469,0.00031746484,0.0001544609,0.00018544703,0.0032360526,0.001502657,0.0082625495],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9933269,0.0010785484,0.00037538988,0.0012745496,0.0031656234,0.0007789695],"domain_scores_gemma":[0.9792763,0.006916662,0.002222127,0.0069790925,0.0041251974,0.00048064318],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031837213,0.0008244562,0.000984853,0.005239545,0.0021625236,0.0044780695,0.0020990376,0.0027290469,0.0050455313],"category_scores_gemma":[0.04169274,0.0009316755,0.0010109334,0.004100944,0.001491258,0.008386256,0.005717893,0.0029225147,0.0028846215],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002940123,0.00015532164,0.022137266,0.00035259346,0.000045032426,0.0007405312,0.001721335,0.012364718,0.014047226,0.2846404,0.019720709,0.64378095],"study_design_scores_gemma":[0.00007824621,0.000073120274,0.0074459836,0.00017715969,0.000083792125,0.0012468136,0.001164369,0.28721738,0.028738553,0.5423686,0.13128076,0.00012522157],"about_ca_topic_score_codex":0.00594134,"about_ca_topic_score_gemma":0.004367982,"teacher_disagreement_score":0.00594134,"about_ca_system_score_codex":0.0012498088,"about_ca_system_score_gemma":0.002539201,"threshold_uncertainty_score":0.016878963},"labels":[],"label_agreement":null},{"id":"W3013047839","doi":"10.1109/iwsc50091.2020.9047642","title":"Clone Swarm: A Cloud Based Code-Clone Analysis Tool","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Swarm behaviour; Computer science; Source code; Software maintenance; Java; Software; Cloud computing; Code (set theory); Fragment (logic); Software system; Operating system; Programming language; Biology; Artificial intelligence; Set (abstract data type); Genetics; Gene","score_opus":0.028259094032296,"score_gpt":0.26298545009374513,"score_spread":0.23472635606144912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3013047839","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026419073,0.0010779666,0.30396083,0.0006766205,0.0002273244,0.000674887,0.019768132,0.63396484,0.013230354],"genre_scores_gemma":[0.2549385,0.0015263499,0.55982774,0.0011721824,0.00020975136,0.0015418094,0.07828366,0.08262933,0.01987073],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998307,0.0001681905,0.00014171652,0.00041773636,0.00083271094,0.00013261032],"domain_scores_gemma":[0.99393,0.002355002,0.00094808114,0.0011724357,0.0011850416,0.00040935804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015708167,0.0015856804,0.0008834856,0.0049701524,0.0009200274,0.0023421638,0.0022543557,0.0010757293,0.008569412],"category_scores_gemma":[0.01256483,0.0009877885,0.0014840751,0.0034174286,0.0007362917,0.004146852,0.003243183,0.0014801641,0.0055057458],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012886482,0.0003439897,0.038406514,0.0020269048,0.0004439648,0.0018778522,0.003089473,0.01254742,0.03171696,0.018401476,0.3910252,0.4988317],"study_design_scores_gemma":[0.0007327135,0.0005275876,0.03751701,0.0007722416,0.000353789,0.0025247654,0.0012553702,0.3925666,0.06314927,0.045201663,0.45491555,0.00048349155],"about_ca_topic_score_codex":0.006226182,"about_ca_topic_score_gemma":0.0067080213,"teacher_disagreement_score":0.008569412,"about_ca_system_score_codex":0.0013317867,"about_ca_system_score_gemma":0.0020211567,"threshold_uncertainty_score":0.02866757},"labels":[],"label_agreement":null},{"id":"W3013132080","doi":"10.1109/iwsc50091.2020.9047643","title":"SemanticCloneBench: A Semantic Code Clone Benchmark using Crowd-Source Knowledge","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Java; Benchmark (surveying); Programming language; clone (Java method); Python (programming language); Source code; Code (set theory)","score_opus":0.049751312979047654,"score_gpt":0.29086922157050293,"score_spread":0.24111790859145527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3013132080","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5541143,0.0010999608,0.34384823,0.0005648478,0.00028077897,0.0022586468,0.017150024,0.06437067,0.016312556],"genre_scores_gemma":[0.5828739,0.000294773,0.35704577,0.0002513691,0.000056269517,0.0017367095,0.049090724,0.004677354,0.003973095],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9914874,0.0022168849,0.0007661703,0.0016614392,0.0034473194,0.00042085032],"domain_scores_gemma":[0.9784468,0.008687181,0.001752367,0.00440685,0.0057845297,0.0009222145],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00532375,0.0017931209,0.0009011672,0.007068262,0.0010284695,0.0017421494,0.0026580151,0.0019246845,0.0016708703],"category_scores_gemma":[0.028275674,0.00047997572,0.0010106165,0.0039495295,0.0012403392,0.0031948162,0.0032491817,0.0011999443,0.0010568948],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023031377,0.0027975624,0.098804176,0.004069133,0.0009719004,0.002695098,0.008399706,0.10603573,0.08095278,0.019832293,0.077258304,0.59588015],"study_design_scores_gemma":[0.00053895556,0.0019043251,0.0638722,0.0005284213,0.00031584335,0.0019696378,0.0034018399,0.68573403,0.13051985,0.022190483,0.08858606,0.00043836792],"about_ca_topic_score_codex":0.006465557,"about_ca_topic_score_gemma":0.0073379097,"teacher_disagreement_score":0.007068262,"about_ca_system_score_codex":0.0013195443,"about_ca_system_score_gemma":0.0020557705,"threshold_uncertainty_score":0.028155029},"labels":[],"label_agreement":null},{"id":"W3013524315","doi":"10.5383/juspn.08.01.001","title":"Software Quality Assessment Algorithm Based on Fuzzy Logic","year":2017,"lang":"en","type":"article","venue":"Journal of Ubiquitous Systems and Pervasive Networks","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Software quality; Fuzzy logic; Data mining; Verification and validation; Software metric; Software; Quality (philosophy); Software quality control; Metric (unit); Software measurement; Algorithm; Software development; Software engineering; Artificial intelligence; Programming language; Mathematics; Statistics; Engineering","score_opus":0.04507118796917641,"score_gpt":0.3327389602104978,"score_spread":0.2876677722413214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3013524315","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059524253,0.00010458125,0.9918994,0.000056184224,0.000013720704,0.0000732239,0.000028565768,0.00023521202,0.0016366376],"genre_scores_gemma":[0.25326577,0.00020538019,0.74438995,0.00007811471,0.00002961794,0.0002748201,0.00022982103,0.00005761784,0.0014688795],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968251,0.00055092265,0.00027877252,0.00061477115,0.0015539923,0.00017655424],"domain_scores_gemma":[0.99802643,0.00078685523,0.00020430377,0.00012160251,0.000810922,0.000049904094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029753565,0.00072476495,0.0008655867,0.0030230356,0.00069936016,0.0017412961,0.0014034209,0.00093744544,0.0021732717],"category_scores_gemma":[0.00632231,0.00023936028,0.0012729125,0.001369274,0.0005439913,0.0018354938,0.00087169866,0.0009734065,0.00034952606],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016313494,0.00016981269,0.0044647213,0.0003297057,0.00013901351,0.0001375543,0.00036625814,0.29387873,0.009493076,0.06259432,0.002225204,0.62603843],"study_design_scores_gemma":[0.00002549744,0.000101034355,0.0008035018,0.00004876997,0.00004651205,0.00008374964,0.000047141082,0.9742371,0.0033694326,0.01917813,0.0020400747,0.000019065801],"about_ca_topic_score_codex":0.0053809653,"about_ca_topic_score_gemma":0.003436851,"teacher_disagreement_score":0.0053809653,"about_ca_system_score_codex":0.0022246481,"about_ca_system_score_gemma":0.0016109077,"threshold_uncertainty_score":0.016141057},"labels":[],"label_agreement":null},{"id":"W3013853519","doi":"10.1109/iwsc50091.2020.9047639","title":"Evaluating Performance of Clone Detection Tools in Detecting Cloned Cochange Candidates","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Computer science; Cloning (programming); Java; Codebase; Software maintenance; Code refactoring; Computational biology; Software; Programming language; Biology; Software system; Genetics; Gene","score_opus":0.09441482906628972,"score_gpt":0.3313859175855781,"score_spread":0.2369710885192884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3013853519","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9533476,0.0014054079,0.03873468,0.000088876586,0.00008102537,0.00016279417,0.0004206643,0.004237344,0.0015215683],"genre_scores_gemma":[0.9078691,0.00041058948,0.08879401,0.000052027637,0.000026128178,0.000087133056,0.0013435282,0.00022387409,0.0011935863],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9923723,0.001955892,0.0010698108,0.0015303207,0.0024402945,0.00063142926],"domain_scores_gemma":[0.91708505,0.06427287,0.0039338404,0.0036985867,0.009266014,0.0017436497],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.00657874,0.0016364729,0.0012303654,0.0073210937,0.0006756303,0.0020098542,0.0014898373,0.0020263821,0.0007760057],"category_scores_gemma":[0.037506994,0.00043729082,0.001047953,0.0034552272,0.0006655745,0.0029360238,0.0013074129,0.00097232347,0.00054389617],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0065036393,0.0023882398,0.2272514,0.0017517208,0.0014960845,0.0010232017,0.001960912,0.079411946,0.08286176,0.0017646265,0.0040314253,0.5895551],"study_design_scores_gemma":[0.00021298057,0.004263261,0.09046469,0.00012738943,0.00067036937,0.0010655541,0.00094781746,0.783553,0.11419912,0.0009459717,0.0033223832,0.00022741719],"about_ca_topic_score_codex":0.0052699083,"about_ca_topic_score_gemma":0.0039728605,"teacher_disagreement_score":0.99342126,"about_ca_system_score_codex":0.0007781097,"about_ca_system_score_gemma":0.0011180408,"threshold_uncertainty_score":0.034792125},"labels":[],"label_agreement":null},{"id":"W3014111015","doi":"10.1109/saner48275.2020.9054846","title":"Associating Code Clones with Association Rules for Change Impact Analysis","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Commit; Programmer; Association rule learning; clone (Java method); Set (abstract data type); Correctness; Precision and recall; Code (set theory); Source code; Data mining; Association (psychology); Change impact analysis; Cloning (programming); Code refactoring; Programming language; Information retrieval; Software; Database; Biology; Gene","score_opus":0.05241528672329757,"score_gpt":0.3179658369072199,"score_spread":0.2655505501839223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014111015","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33907914,0.0015116794,0.64299893,0.0004318537,0.000115829775,0.000587987,0.0035724065,0.009236781,0.0024653906],"genre_scores_gemma":[0.5999411,0.00037118676,0.39466697,0.000121035824,0.000056332046,0.00031201157,0.003482607,0.0001911568,0.0008575636],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9938775,0.0007838458,0.0009113288,0.0015404104,0.0025590346,0.0003278774],"domain_scores_gemma":[0.957258,0.024629207,0.0077560293,0.004153115,0.0055308603,0.0006728261],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004410566,0.0010330677,0.0013124101,0.013864561,0.0009166141,0.0021801766,0.001715917,0.0011599349,0.000980273],"category_scores_gemma":[0.029990086,0.0005939677,0.0015441335,0.008572406,0.00073413097,0.002769743,0.001617794,0.0016535777,0.0006412819],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037568455,0.00060363807,0.3879381,0.0004876523,0.0005794373,0.0014719366,0.0011427519,0.042378422,0.011964384,0.006224885,0.0037934165,0.54303974],"study_design_scores_gemma":[0.000052007435,0.0003026001,0.07675285,0.00021761362,0.0004925404,0.0025243142,0.0007180877,0.869045,0.020353245,0.019660762,0.009744348,0.00013671211],"about_ca_topic_score_codex":0.006384314,"about_ca_topic_score_gemma":0.008846749,"teacher_disagreement_score":0.013864561,"about_ca_system_score_codex":0.00082045497,"about_ca_system_score_gemma":0.001571927,"threshold_uncertainty_score":0.023325562},"labels":[],"label_agreement":null},{"id":"W3014294688","doi":"10.1145/3341105.3374008","title":"Detecting architectural integrity violation patterns using machine learning","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada; Ontario Centres of Excellence","keywords":"Maintainability; Computer science; Architectural pattern; Set (abstract data type); Software engineering; Artificial intelligence; Machine learning; Software quality; Software; Software system; Software development; Programming language; Software construction","score_opus":0.050813695279742645,"score_gpt":0.28644156137671245,"score_spread":0.2356278660969698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014294688","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73211515,0.00078528107,0.25829175,0.0008877761,0.00007306363,0.00018758934,0.0012501847,0.0044756965,0.0019335524],"genre_scores_gemma":[0.93330795,0.00011657683,0.064166844,0.00008331958,0.000026689448,0.00006716071,0.0016711032,0.000047239537,0.0005131764],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99744886,0.00060522556,0.00027373747,0.0006791,0.0007046427,0.00028839463],"domain_scores_gemma":[0.9858979,0.007211936,0.003126796,0.001492026,0.0019272204,0.00034411572],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021423367,0.00096847076,0.0008182283,0.005244309,0.0005560269,0.0014240661,0.0016493504,0.0014062177,0.00066188036],"category_scores_gemma":[0.0125398245,0.0003912508,0.00077106466,0.0029702855,0.00065847026,0.0015899385,0.0009580635,0.0014625667,0.0004239468],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003087088,0.0010562508,0.32761753,0.00026042934,0.00034048315,0.00082200754,0.00050527643,0.16915116,0.012684652,0.0017792179,0.0053259204,0.48014843],"study_design_scores_gemma":[0.000010585815,0.000068642985,0.014012966,0.000016361364,0.000024709303,0.00012113577,0.000092138536,0.97918177,0.0032835829,0.0026384946,0.00053177227,0.000017874583],"about_ca_topic_score_codex":0.005434436,"about_ca_topic_score_gemma":0.007091935,"teacher_disagreement_score":0.005434436,"about_ca_system_score_codex":0.00086172874,"about_ca_system_score_gemma":0.0010369355,"threshold_uncertainty_score":0.011329889},"labels":[],"label_agreement":null},{"id":"W3014514721","doi":"10.1109/saner48275.2020.9054869","title":"HistoRank: History-Based Ranking of Co-change Candidates","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Ranking (information retrieval); Computer science; Rank (graph theory); False positive paradox; Machine learning; Ranking SVM; Data mining; Association rule learning; Association (psychology); Artificial intelligence; Information retrieval; Mathematics; Psychology","score_opus":0.06699375120327251,"score_gpt":0.2664511262385754,"score_spread":0.19945737503530292,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014514721","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64378715,0.007776782,0.259538,0.0010298868,0.0011138589,0.0013805514,0.016516784,0.056729846,0.012127108],"genre_scores_gemma":[0.7854863,0.0007768693,0.18849915,0.00011919564,0.00019744126,0.00033774978,0.017144712,0.00071803725,0.0067204675],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965204,0.0005560516,0.00037423594,0.0007081197,0.0015938355,0.00024740625],"domain_scores_gemma":[0.9855385,0.005322529,0.0016263562,0.0022880824,0.004399126,0.0008254416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003825554,0.0016037083,0.0013676956,0.009211433,0.0010246221,0.0022738727,0.001965013,0.0007516268,0.0030236002],"category_scores_gemma":[0.019705197,0.00050174043,0.00065175135,0.0045435056,0.00036589964,0.0030027882,0.0010952366,0.00072705897,0.001455207],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017545056,0.0007082539,0.10861664,0.0011793424,0.0005194041,0.00038675658,0.000566467,0.029839726,0.012391634,0.0031855335,0.055070516,0.78578126],"study_design_scores_gemma":[0.0004027653,0.0017145365,0.066979244,0.00013030387,0.0003810375,0.0010374297,0.0007702691,0.86041147,0.03115456,0.006583531,0.030120363,0.00031452943],"about_ca_topic_score_codex":0.008083825,"about_ca_topic_score_gemma":0.017213728,"teacher_disagreement_score":0.009211433,"about_ca_system_score_codex":0.00097552786,"about_ca_system_score_gemma":0.0023459843,"threshold_uncertainty_score":0.020231724},"labels":[],"label_agreement":null},{"id":"W3014646158","doi":"10.1109/saner48275.2020.9054834","title":"We Are Family: Analyzing Communication in GitHub Software Repositories and Their Forks","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Fork (system call); Computer science; Software; Process (computing); World Wide Web; Operating system","score_opus":0.028122138691581852,"score_gpt":0.25106443437547954,"score_spread":0.22294229568389767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014646158","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99755484,0.0000959305,0.0009623165,0.00007604412,0.0000022858108,0.000025008338,0.00026216128,0.00006499903,0.00095651107],"genre_scores_gemma":[0.9967609,0.00008534824,0.0017699753,0.000019885805,0.0000074569216,0.00005173741,0.00068887166,0.000033555334,0.0005823971],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9946517,0.0022632421,0.0005179968,0.0006499478,0.0013039213,0.0006131164],"domain_scores_gemma":[0.94669354,0.028008875,0.0127692865,0.0036309974,0.0062136548,0.0026836358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042677056,0.0003581236,0.00037728384,0.0064421785,0.0015250355,0.0018686943,0.00070542406,0.0006167073,0.0017020978],"category_scores_gemma":[0.036602367,0.0002587151,0.00044011933,0.0075585856,0.0010419665,0.00473884,0.0028129434,0.00062462146,0.0006136593],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025148687,0.00011014406,0.9218711,0.000115024486,0.00006665099,0.00046389873,0.03641129,0.0007488917,0.00092696433,0.0011181986,0.0015352389,0.036381043],"study_design_scores_gemma":[0.000017415477,0.00019204195,0.9289768,0.00006936076,0.000053256663,0.0007567313,0.054146912,0.008084328,0.0010405201,0.0013799948,0.005230497,0.000052199495],"about_ca_topic_score_codex":0.01402349,"about_ca_topic_score_gemma":0.011121043,"teacher_disagreement_score":0.01402349,"about_ca_system_score_codex":0.0016089336,"about_ca_system_score_gemma":0.0015637818,"threshold_uncertainty_score":0.027883708},"labels":[],"label_agreement":null},{"id":"W3014728464","doi":"10.1109/saner48275.2020.9054848","title":"Studying Developer Reading Behavior on Stack Overflow during API Summarization Tasks","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Carleton University","funders":"","keywords":"Computer science; Automatic summarization; Source code; World Wide Web; Reading (process); Software; Program comprehension; Gaze; Reading comprehension; Code (set theory); Information retrieval; Artificial intelligence; Programming language; Software system","score_opus":0.04476517789639843,"score_gpt":0.2741167396505266,"score_spread":0.22935156175412813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014728464","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9993618,0.000016659424,0.00021467856,0.0000106682055,0.000001229469,0.000011923044,0.000047264923,0.000019337365,0.00031640983],"genre_scores_gemma":[0.9974771,0.00006447385,0.0010551439,0.000025136604,0.000004136079,0.00004783829,0.00023734184,0.000019379415,0.0010693944],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99939704,0.00020452932,0.000051086074,0.00015808389,0.000119866105,0.000069344074],"domain_scores_gemma":[0.98754054,0.008121023,0.0018138876,0.00055020547,0.0014702363,0.0005040672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00073297817,0.00032323858,0.0002817712,0.0011254702,0.0003917093,0.0005409131,0.00026453883,0.00045374528,0.0013650417],"category_scores_gemma":[0.014803074,0.0002699827,0.00017120596,0.00048462546,0.00029152734,0.00051717187,0.0004897836,0.0003795289,0.00040576464],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088494737,0.0007016705,0.7146591,0.000299308,0.00010643118,0.0008818109,0.08228213,0.00059651886,0.119704664,0.00018951291,0.001424877,0.078268945],"study_design_scores_gemma":[0.000014299998,0.00063334854,0.9809823,0.000026641823,0.00002859011,0.00022486813,0.009848291,0.0016765828,0.0054867207,0.00007442804,0.0009639483,0.000039995753],"about_ca_topic_score_codex":0.005952872,"about_ca_topic_score_gemma":0.0129093295,"teacher_disagreement_score":0.005952872,"about_ca_system_score_codex":0.0003701374,"about_ca_system_score_gemma":0.00022598596,"threshold_uncertainty_score":0.01183641},"labels":[],"label_agreement":null},{"id":"W3014903559","doi":"10.1145/1449955.1449790","title":"Enabling static analysis for partial java programs","year":2008,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Programming language; Java; Source code; Static analysis; Partial evaluation; Class (philosophy); Type inference; Theoretical computer science; Inference; Artificial intelligence","score_opus":0.06896249739321578,"score_gpt":0.30438009191292875,"score_spread":0.23541759451971297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014903559","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0618417,0.00012385762,0.89656365,0.00014555687,0.000021853913,0.000091977185,0.00038986656,0.039461814,0.0013596835],"genre_scores_gemma":[0.41235104,0.00017316152,0.57998985,0.000111074565,0.00004453399,0.00015314599,0.0015226601,0.0040978463,0.001556664],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99561596,0.0011499015,0.0003200774,0.00075311115,0.0017064195,0.0004544727],"domain_scores_gemma":[0.9756378,0.014059959,0.00180728,0.005664486,0.0025090235,0.00032147244],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038044762,0.0010142608,0.0013378151,0.0029678363,0.0013495653,0.0025734091,0.0021662186,0.0010511121,0.0024963822],"category_scores_gemma":[0.026606012,0.0012588222,0.0013358595,0.002399681,0.0019466386,0.0039037236,0.002989615,0.0015708406,0.001040316],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008047699,0.00036629476,0.026127458,0.00086889905,0.0001926196,0.0009156134,0.0036161446,0.12016545,0.09434258,0.056060567,0.010167905,0.6863717],"study_design_scores_gemma":[0.00005198115,0.00013547162,0.0045821574,0.00012457269,0.00012472521,0.00046528951,0.00039034398,0.84680563,0.08143999,0.05121037,0.014549745,0.00011981501],"about_ca_topic_score_codex":0.00754783,"about_ca_topic_score_gemma":0.011365341,"teacher_disagreement_score":0.00754783,"about_ca_system_score_codex":0.0010500976,"about_ca_system_score_gemma":0.0032527386,"threshold_uncertainty_score":0.020120203},"labels":[],"label_agreement":null},{"id":"W3014978656","doi":"","title":"GitHub Repositories with Links to Academic Papers: Open Access, Traceability, and Evolution","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Traceability; Computer science; Source code; World Wide Web; Data science; Open source; Open source software; Code (set theory); Software; Software engineering; Programming language","score_opus":0.11488166432844935,"score_gpt":0.26148531083092685,"score_spread":0.1466036465024775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014978656","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84104854,0.027098201,0.036411468,0.009517856,0.0011441819,0.00065682945,0.034814496,0.006802838,0.042505596],"genre_scores_gemma":[0.8793764,0.008469757,0.041550167,0.0016465756,0.0010185497,0.0007310747,0.047591615,0.0036226183,0.015993401],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9650163,0.00936161,0.0048360773,0.003809605,0.015432058,0.0015443204],"domain_scores_gemma":[0.6050887,0.19740856,0.09577524,0.0492072,0.044884656,0.007635597],"candidate_categories":["scholarly_communication","open_science"],"consensus_categories":[],"category_scores_codex":[0.025860837,0.00078669155,0.001086722,0.056404803,0.0033597192,0.014767824,0.002254109,0.0017151346,0.0062530646],"category_scores_gemma":[0.26812106,0.00072810624,0.00090593065,0.085825935,0.0022725393,0.012133634,0.009933225,0.0017009006,0.0028797511],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007840253,0.00035800494,0.5517058,0.005934572,0.0009336424,0.0035209092,0.0276835,0.0016309136,0.004210542,0.029935671,0.04885742,0.32444492],"study_design_scores_gemma":[0.00015178426,0.00024903216,0.61451685,0.0042627226,0.0008609594,0.00451169,0.014720604,0.0064991387,0.0074779014,0.032287348,0.3141265,0.00033545346],"about_ca_topic_score_codex":0.00422678,"about_ca_topic_score_gemma":0.005004787,"teacher_disagreement_score":0.9977459,"about_ca_system_score_codex":0.0021078975,"about_ca_system_score_gemma":0.0037947495,"threshold_uncertainty_score":0.13676679},"labels":[],"label_agreement":null},{"id":"W3016122787","doi":"10.1088/1742-6596/1487/1/012017","title":"A Systematic Study for Learning-Based Software Defect Prediction","year":2020,"lang":"en","type":"article","venue":"Journal of Physics Conference Series","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Artificial intelligence; Machine learning; Deep learning; Process (computing); Field (mathematics); Software bug; Software; Software development; Software engineering; Programming language","score_opus":0.03477111679928294,"score_gpt":0.264824814736505,"score_spread":0.23005369793722205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3016122787","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2757303,0.018679755,0.69374895,0.0019506683,0.00031179466,0.000719477,0.0005918482,0.00056739885,0.007699761],"genre_scores_gemma":[0.8749819,0.0047855964,0.117462344,0.00026317468,0.00017109855,0.0002834981,0.00075297296,0.000050511127,0.001248925],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961979,0.0017275121,0.0003491378,0.00070600345,0.0008957959,0.00012366162],"domain_scores_gemma":[0.9801702,0.011578267,0.001203862,0.0014551881,0.005383652,0.00020889385],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055901026,0.0007454374,0.0006567001,0.00228203,0.00048825116,0.0013778526,0.0010321612,0.00072287605,0.0007766237],"category_scores_gemma":[0.023454685,0.00033222808,0.00093636307,0.0018011482,0.0007204916,0.0027714863,0.0010452205,0.0012003747,0.00022665253],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024780908,0.001023325,0.08712153,0.0025800737,0.0006683741,0.00042825902,0.0006742139,0.11642353,0.0063212207,0.035148844,0.005592867,0.74376994],"study_design_scores_gemma":[0.000025807807,0.00052473415,0.01263228,0.00059538585,0.000292845,0.00021623475,0.0003329208,0.9560578,0.0055978685,0.01852565,0.005159015,0.000039491122],"about_ca_topic_score_codex":0.004172017,"about_ca_topic_score_gemma":0.002890778,"teacher_disagreement_score":0.0055901026,"about_ca_system_score_codex":0.0012375219,"about_ca_system_score_gemma":0.0018976869,"threshold_uncertainty_score":0.029563606},"labels":[],"label_agreement":null},{"id":"W3016705130","doi":"10.1371/journal.pone.0231731","title":"Empirical study of the relationship between design patterns and code smells","year":2020,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code smell; Computer science; Software design pattern; Software design; Object-oriented design; Association (psychology); Software; Code (set theory); Artificial intelligence; Software development; Programming language; Software quality; Psychology; Set (abstract data type)","score_opus":0.3117284591746598,"score_gpt":0.32797044803361974,"score_spread":0.016241988858959944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3016705130","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.998288,0.00006278086,0.0011585836,0.000033011205,0.0000022359463,0.000024515242,0.00010387125,0.000010685573,0.0003163452],"genre_scores_gemma":[0.9982651,0.00003186052,0.0013457447,0.000008430877,0.0000036553893,0.000035022847,0.00020352993,0.0000060819793,0.00010053405],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9880691,0.0043390724,0.001982274,0.0017788727,0.0033835731,0.0004471742],"domain_scores_gemma":[0.5775101,0.3342391,0.05928107,0.01121596,0.015020426,0.0027333053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0096813515,0.00040168068,0.0003594852,0.0030753873,0.0004599855,0.0011349975,0.00063855853,0.0007575001,0.001087974],"category_scores_gemma":[0.11892786,0.00033860028,0.0005293902,0.0030043928,0.0011944631,0.0020014131,0.0013132772,0.0010958314,0.00019388252],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013103746,0.0002104501,0.9881131,0.000096097894,0.00009993486,0.00009313532,0.0013185529,0.00048146883,0.0006533105,0.00012732614,0.000084943946,0.008590636],"study_design_scores_gemma":[0.000014142614,0.00039350678,0.9897928,0.000033885077,0.0000598621,0.0003109916,0.0018679214,0.0056235883,0.0011195777,0.00029991055,0.000465203,0.000018528557],"about_ca_topic_score_codex":0.0010231956,"about_ca_topic_score_gemma":0.0016880931,"teacher_disagreement_score":0.0096813515,"about_ca_system_score_codex":0.000554334,"about_ca_system_score_gemma":0.0005537191,"threshold_uncertainty_score":0.05120051},"labels":[],"label_agreement":null},{"id":"W3017077459","doi":"10.1007/978-3-030-44914-8_6","title":"Concise Read-Only Specifications for Better Synthesis of Programs with Pointers","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Yale-NUS College; Yale University; National Science Foundation","keywords":"Computer science; Program synthesis; Correctness; Heap (data structure); Programming language; Separation logic; High-level synthesis; Intuition; Theoretical computer science; Embedded system","score_opus":0.034887188510758595,"score_gpt":0.2531256888930704,"score_spread":0.21823850038231182,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3017077459","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055052765,0.00017996848,0.9335807,0.00024038226,0.00005343981,0.000081348546,0.00021249522,0.006819004,0.0037799303],"genre_scores_gemma":[0.45419902,0.00021627225,0.5388739,0.00029027055,0.00003020203,0.00018675458,0.000628463,0.0021962014,0.0033788935],"study_design_codex":"bench_or_experimental","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9970854,0.0010247035,0.00028822714,0.00027523367,0.0011496465,0.0001768319],"domain_scores_gemma":[0.98761266,0.007469098,0.0008221926,0.0028118112,0.0011366705,0.0001474408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003298145,0.0007166752,0.0005019265,0.0007424845,0.00039264883,0.0015628255,0.0013628976,0.0007877545,0.0062557673],"category_scores_gemma":[0.011284658,0.00054256106,0.000748556,0.0005466862,0.0014316167,0.0030983977,0.0016143065,0.0017776486,0.0010088332],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067188725,0.00038342315,0.003388689,0.0018563826,0.000108248525,0.0004893508,0.0010034203,0.2274225,0.3034917,0.2304153,0.005317324,0.22545184],"study_design_scores_gemma":[0.0001102765,0.0004176672,0.00058708974,0.00018745645,0.00015396942,0.0002633295,0.00020206961,0.4353301,0.4379295,0.093171805,0.0315706,0.000076165015],"about_ca_topic_score_codex":0.00054934877,"about_ca_topic_score_gemma":0.0010822719,"teacher_disagreement_score":0.0062557673,"about_ca_system_score_codex":0.0007063891,"about_ca_system_score_gemma":0.001236767,"threshold_uncertainty_score":0.020927668},"labels":[],"label_agreement":null},{"id":"W3018272941","doi":"10.1016/j.scico.2020.102472","title":"The prevalence and severity of persistent ambiguity in software requirements specifications: Is a special effort needed to find them?","year":2020,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Sheridan College","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Ambiguity; Software requirements specification; Sampling (signal processing); Process (computing); System requirements specification; Software engineering; Software; Software development; Risk analysis (engineering); Programming language; Software design","score_opus":0.06871544950554817,"score_gpt":0.2913727760571159,"score_spread":0.22265732655156775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3018272941","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.813342,0.010297384,0.11787476,0.042731676,0.00070791,0.00015438815,0.0003993799,0.0006103154,0.013882212],"genre_scores_gemma":[0.9674169,0.0015298158,0.02831609,0.0017588069,0.00032408372,0.00004319464,0.00019726454,0.00013393338,0.00027995746],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9165481,0.036770564,0.012110358,0.004285168,0.028103301,0.002182605],"domain_scores_gemma":[0.41051528,0.4296969,0.08734978,0.03215988,0.03623209,0.004046003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0722291,0.0004737972,0.0008223739,0.0063972254,0.0018948077,0.0063712336,0.0025487791,0.0035983128,0.0016332516],"category_scores_gemma":[0.4492901,0.0011156214,0.0008722159,0.005146103,0.0056214384,0.023203164,0.006331486,0.004919088,0.0004708859],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008107107,0.00037059106,0.38941896,0.0015566647,0.00036531663,0.0017815772,0.037062164,0.0030269846,0.007530917,0.06768552,0.0057846108,0.4846059],"study_design_scores_gemma":[0.00023590123,0.0010200934,0.29500917,0.00526282,0.0009385996,0.024969136,0.12770814,0.03200998,0.01037212,0.45919865,0.04250785,0.00076755154],"about_ca_topic_score_codex":0.001683704,"about_ca_topic_score_gemma":0.0018449523,"teacher_disagreement_score":0.0722291,"about_ca_system_score_codex":0.0013508346,"about_ca_system_score_gemma":0.0035997953,"threshold_uncertainty_score":0.38198858},"labels":[],"label_agreement":null},{"id":"W3018447383","doi":"10.1007/s10664-020-09819-6","title":"What do Programmers Discuss about Deep Learning Frameworks","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":89,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Deep learning; Workflow; Computer science; Latent Dirichlet allocation; Leverage (statistics); Artificial intelligence; Topic model; Data science; Profiling (computer programming); World Wide Web; Machine learning","score_opus":0.01947834175930554,"score_gpt":0.2802885285859658,"score_spread":0.2608101868266603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3018447383","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03961806,0.025653223,0.16569036,0.6940899,0.005777477,0.000049728835,0.00029004982,0.000631028,0.06820028],"genre_scores_gemma":[0.69162405,0.03574362,0.097078905,0.13169253,0.01473747,0.00019993496,0.00060862734,0.0019149282,0.02639988],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9898347,0.004696473,0.00047088417,0.0011224025,0.0027704025,0.0011050804],"domain_scores_gemma":[0.9206953,0.056469396,0.004114738,0.0050578797,0.010067845,0.0035948043],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01687423,0.00065518566,0.0006053128,0.0022965523,0.0033350934,0.010046354,0.0021529181,0.005931432,0.007407367],"category_scores_gemma":[0.11264685,0.00073939044,0.00073387666,0.0028980966,0.0067638187,0.031993896,0.0025876875,0.008907746,0.0015571446],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001883414,0.00013982429,0.011869071,0.0011793923,0.00013071358,0.00027281343,0.0061049946,0.0028189027,0.0013386308,0.5466093,0.090537645,0.33881035],"study_design_scores_gemma":[0.000039018578,0.00006488232,0.0035754961,0.002943536,0.000108349705,0.00079435884,0.01025494,0.0052162986,0.0028133434,0.6622452,0.31184426,0.00010026101],"about_ca_topic_score_codex":0.0028736605,"about_ca_topic_score_gemma":0.003119077,"teacher_disagreement_score":0.01687423,"about_ca_system_score_codex":0.0022939001,"about_ca_system_score_gemma":0.003750404,"threshold_uncertainty_score":0.08924049},"labels":[],"label_agreement":null},{"id":"W3019290077","doi":"10.1002/spe.2830","title":"Vocabulary and time based bug‐assignment: A recommender system for open‐source projects","year":2020,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Vocabulary; Relevance (law); Key (lock); Intuition; Ranking (information retrieval); Task (project management); Information retrieval; World Wide Web; Software engineering; Data science; Artificial intelligence; Computer security","score_opus":0.04127924435912676,"score_gpt":0.30250029074955476,"score_spread":0.261221046390428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3019290077","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46547922,0.015241418,0.43850055,0.0035195302,0.0012315675,0.0015142018,0.023667471,0.03746926,0.013376682],"genre_scores_gemma":[0.599873,0.0021966903,0.35909012,0.0006133405,0.00049717125,0.0005473901,0.028910179,0.00050646893,0.0077656573],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969317,0.0006934058,0.00036445478,0.0009926258,0.00084199046,0.0001758126],"domain_scores_gemma":[0.98926896,0.0050668507,0.0011093442,0.0011508905,0.00279148,0.0006123673],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037356704,0.0013431695,0.0016577803,0.009380295,0.0009826636,0.0016879974,0.0022420548,0.0016806233,0.0020380362],"category_scores_gemma":[0.019551031,0.0006572977,0.001159089,0.0049700937,0.0002661785,0.003038537,0.0015926076,0.0014582576,0.0027713478],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007232324,0.0009130163,0.10070902,0.0015093135,0.00065538025,0.0007777232,0.0012727543,0.01678901,0.012032712,0.0028458377,0.09094586,0.7708262],"study_design_scores_gemma":[0.0003275654,0.000855008,0.057354107,0.00032800602,0.0006412792,0.0012603349,0.00083571445,0.8726486,0.0071745287,0.007798792,0.050498985,0.00027703814],"about_ca_topic_score_codex":0.024505695,"about_ca_topic_score_gemma":0.045139927,"teacher_disagreement_score":0.024505695,"about_ca_system_score_codex":0.0009719598,"about_ca_system_score_gemma":0.0014460628,"threshold_uncertainty_score":0.04872614},"labels":[],"label_agreement":null},{"id":"W3020617474","doi":"10.1016/j.jss.2020.110610","title":"Code smells and refactoring: A tertiary systematic review of challenges and observations","year":2020,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":206,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Université du Québec à Chicoutimi","funders":"","keywords":"Code refactoring; Code smell; Code (set theory); Systematic review; Quality (philosophy); Code review; Software","score_opus":0.07286601472112181,"score_gpt":0.27273289922332516,"score_spread":0.19986688450220336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3020617474","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021240782,0.9698057,0.0018783371,0.002468693,0.00031695698,0.00095509656,0.0023644967,0.00004775396,0.0009221798],"genre_scores_gemma":[0.18147065,0.80047804,0.0082627535,0.0045338767,0.00027907256,0.0019010932,0.0025262346,0.000099548044,0.00044868432],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.944409,0.018442215,0.01952694,0.004837093,0.01152627,0.0012584614],"domain_scores_gemma":[0.7086046,0.2045894,0.04264225,0.010629807,0.031243797,0.0022901217],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.057551235,0.0012329057,0.00522261,0.021312214,0.0011129339,0.004158763,0.0025284006,0.0020225085,0.0020658327],"category_scores_gemma":[0.23582605,0.0013915051,0.006425898,0.015973888,0.002187781,0.0047405986,0.0051051946,0.0019814663,0.0003459204],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070559763,0.00011931511,0.019790096,0.7208341,0.013453864,0.00041920252,0.007902251,0.00031314688,0.0013375251,0.0012900666,0.0072331717,0.22660163],"study_design_scores_gemma":[0.00023425925,0.00059164106,0.032834817,0.8777639,0.035844244,0.0007895332,0.0060398886,0.0002677977,0.0010955783,0.0015944781,0.042805288,0.00013856044],"about_ca_topic_score_codex":0.0104884235,"about_ca_topic_score_gemma":0.04090222,"teacher_disagreement_score":0.94244874,"about_ca_system_score_codex":0.005602843,"about_ca_system_score_gemma":0.030790756,"threshold_uncertainty_score":0.30436367},"labels":[],"label_agreement":null},{"id":"W3020858756","doi":"","title":"Ranking co-change candidates of micro-clones.","year":2019,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Ranking (information retrieval); Computer science; Information retrieval","score_opus":0.10667489979190946,"score_gpt":0.4009914658971738,"score_spread":0.29431656610526435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3020858756","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8636196,0.008379668,0.101547465,0.001225612,0.0009662716,0.0007507848,0.0070827836,0.004281824,0.012145948],"genre_scores_gemma":[0.85817754,0.0009300917,0.103109114,0.0003633452,0.0003458062,0.0003366684,0.020195246,0.0009328011,0.015609446],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954721,0.0008175891,0.00030104513,0.0009192454,0.0020305233,0.0004595377],"domain_scores_gemma":[0.98126286,0.007943436,0.0017501432,0.0022874288,0.005210816,0.00154533],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021113115,0.0008617415,0.0014046397,0.007965453,0.0017327012,0.0025063762,0.0016235699,0.0020253255,0.006979473],"category_scores_gemma":[0.016833507,0.0003513028,0.0014329269,0.0051402915,0.0005052367,0.0016990798,0.0014694437,0.0007936274,0.0031748791],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029414033,0.0009277748,0.23871002,0.0012812947,0.000810509,0.002558025,0.00081941334,0.013001191,0.05479642,0.0053937812,0.04442275,0.63433754],"study_design_scores_gemma":[0.0007021156,0.0033484052,0.28058103,0.0005573086,0.0023682849,0.0127532,0.0039399206,0.42774865,0.11124034,0.019449811,0.13696225,0.0003486956],"about_ca_topic_score_codex":0.0036854567,"about_ca_topic_score_gemma":0.010611396,"teacher_disagreement_score":0.007965453,"about_ca_system_score_codex":0.0007912521,"about_ca_system_score_gemma":0.0018385173,"threshold_uncertainty_score":0.02334863},"labels":[],"label_agreement":null},{"id":"W3021573492","doi":"10.1007/s10664-021-09956-6","title":"On systematically building a controlled natural language for functional requirements","year":2021,"lang":"en","type":"preprint","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds National de la Recherche Luxembourg","keywords":"Computer science; Domain (mathematical analysis); Natural language; Vagueness; Context (archaeology); Ambiguity; Flexibility (engineering); Quality (philosophy); Popularity; Grammar; Natural language processing; Artificial intelligence; Programming language; Linguistics; Psychology","score_opus":0.034998506852739225,"score_gpt":0.31997698353604714,"score_spread":0.2849784766833079,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3021573492","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009754392,0.00016735354,0.9805519,0.0012521847,0.000036191082,0.0023532836,0.00071886176,0.0006756776,0.00449007],"genre_scores_gemma":[0.03733023,0.00014787748,0.9580964,0.00029986753,0.000014491738,0.0021791845,0.0011997448,0.00017063793,0.00056165253],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9471646,0.037743956,0.0037020212,0.0033544223,0.0074230437,0.0006119171],"domain_scores_gemma":[0.7570104,0.20154248,0.009327099,0.013377129,0.01779956,0.000943272],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04404705,0.0016032204,0.00073986617,0.005168967,0.002259487,0.0046704304,0.0031339612,0.0016218092,0.004979904],"category_scores_gemma":[0.13591447,0.0013615399,0.0020964316,0.002665836,0.0067081195,0.009364893,0.0059496984,0.0035096135,0.0014420211],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019913197,0.0007623284,0.0069116415,0.0051422333,0.00012027992,0.00092636107,0.038829446,0.039228097,0.027528943,0.5508078,0.009264402,0.32027936],"study_design_scores_gemma":[0.0004144345,0.0010709916,0.0042398265,0.006257536,0.00017561455,0.0016940432,0.022972494,0.23294689,0.03899611,0.41210905,0.27860978,0.0005132936],"about_ca_topic_score_codex":0.008099574,"about_ca_topic_score_gemma":0.01650586,"teacher_disagreement_score":0.04404705,"about_ca_system_score_codex":0.00431125,"about_ca_system_score_gemma":0.012778916,"threshold_uncertainty_score":0.2329458},"labels":[],"label_agreement":null},{"id":"W3021870097","doi":"10.1007/978-3-030-47240-5_5","title":"Emotional Contagion in Open Software Collaborations","year":2020,"lang":"en","type":"book-chapter","venue":"IFIP advances in information and communication technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Emotional contagion; Affect (linguistics); Software; Computer science; Production (economics); Open source software; Order (exchange); Locality; Database; Psychology; Business; Social psychology; Communication; Economics; Microeconomics; Operating system; Finance","score_opus":0.018195107358276057,"score_gpt":0.28374619355265485,"score_spread":0.2655510861943788,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3021870097","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3956076,0.0079604965,0.031854857,0.008504855,0.000857261,0.00004926276,0.00003990428,0.0000873501,0.55503845],"genre_scores_gemma":[0.9849407,0.0008753841,0.0012576652,0.00023736263,0.00017863345,0.000032991284,0.000016037162,0.000017788636,0.012443467],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9988418,0.0006453168,0.000027104075,0.00009319858,0.00021540502,0.0001771657],"domain_scores_gemma":[0.99616265,0.0028626577,0.00033015327,0.00014545451,0.00015355734,0.00034557292],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011520155,0.00028259016,0.00023247462,0.0005069009,0.0019819236,0.0037838323,0.00059046276,0.0014427066,0.0057329205],"category_scores_gemma":[0.006330232,0.00017934422,0.00030026832,0.0006994658,0.0028862471,0.0042789113,0.0041011013,0.0019942124,0.0004429237],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002529859,0.00014032058,0.006412944,0.00032741402,0.00006552772,0.0013302674,0.08900587,0.002680862,0.002839872,0.7600523,0.010385948,0.12650566],"study_design_scores_gemma":[0.000044248965,0.00017645149,0.021452878,0.0004836092,0.00007649795,0.0015633735,0.0632738,0.012078274,0.0010897277,0.8183985,0.081287935,0.0000746219],"about_ca_topic_score_codex":0.00064107386,"about_ca_topic_score_gemma":0.0009994903,"teacher_disagreement_score":0.0057329205,"about_ca_system_score_codex":0.0009046579,"about_ca_system_score_gemma":0.0005199044,"threshold_uncertainty_score":0.01917857},"labels":[],"label_agreement":null},{"id":"W3022206766","doi":"","title":"Detection of feature interaction in dynamic scripting languages.","year":2019,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Toronto Metropolitan University","funders":"","keywords":"Computer science; Scripting language; Feature (linguistics); Programming language; Artificial intelligence; Natural language processing; Linguistics","score_opus":0.045746765441935795,"score_gpt":0.38913945363931796,"score_spread":0.34339268819738217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022206766","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6695727,0.00071380317,0.30386627,0.00027898155,0.000115789015,0.0002759633,0.0013083572,0.016089208,0.0077788685],"genre_scores_gemma":[0.9179281,0.000084437284,0.076891184,0.00008558583,0.000026822638,0.000115326926,0.0015657955,0.0007004481,0.0026022857],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9972109,0.00049367984,0.00020357149,0.0005780978,0.0011744328,0.00033936443],"domain_scores_gemma":[0.99085706,0.004266413,0.0014417811,0.0015150751,0.0014158781,0.0005037532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013722909,0.00046941594,0.00055019953,0.0015758038,0.00060534105,0.0012622584,0.0012773813,0.0011114802,0.002079376],"category_scores_gemma":[0.008655443,0.0004688167,0.0005068806,0.0009356336,0.00063320086,0.0016080415,0.0016186436,0.0011501428,0.00074258854],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018300646,0.00072567177,0.11784739,0.00073875603,0.00018713812,0.0031804806,0.0031144687,0.011600342,0.44408682,0.020187946,0.013491617,0.38300928],"study_design_scores_gemma":[0.000102155514,0.00087589415,0.06852911,0.0001458288,0.00015057402,0.0041579464,0.00091819314,0.5922586,0.28106079,0.017753696,0.03391375,0.00013350951],"about_ca_topic_score_codex":0.0015572831,"about_ca_topic_score_gemma":0.0017787395,"teacher_disagreement_score":0.002079376,"about_ca_system_score_codex":0.0004916778,"about_ca_system_score_gemma":0.00074421085,"threshold_uncertainty_score":0.0072574615},"labels":[],"label_agreement":null},{"id":"W3022352710","doi":"","title":"Predicting bug report fields using stack traces and categorical attributes.","year":2019,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Categorical variable; Stack (abstract data type); Computer science; Artificial intelligence; Natural language processing; Programming language; Machine learning","score_opus":0.12320291689675855,"score_gpt":0.4018137944970319,"score_spread":0.27861087760027337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022352710","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94186854,0.0017996576,0.022861628,0.00064930494,0.00017002788,0.00013707424,0.023361359,0.0073926547,0.0017598071],"genre_scores_gemma":[0.95050824,0.00041004852,0.025570067,0.00004453115,0.00007099816,0.000060199352,0.022087133,0.0001566422,0.0010921095],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981281,0.00035128003,0.00022059954,0.0003479844,0.0007744662,0.0001775389],"domain_scores_gemma":[0.97335637,0.012392391,0.0057273423,0.002227771,0.0048559345,0.0014401283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023321696,0.0008740287,0.00055245444,0.010556773,0.00038250536,0.0012881474,0.0008364075,0.0011145486,0.0014762213],"category_scores_gemma":[0.022046678,0.0003290774,0.00079033204,0.0055037094,0.0002801757,0.001886489,0.0009598759,0.0009594328,0.0012527144],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007858345,0.00060108025,0.78650415,0.0003952696,0.0002232117,0.00034594297,0.00025232197,0.0102027105,0.0065192683,0.0008500273,0.015011237,0.17830896],"study_design_scores_gemma":[0.00012897688,0.0014066683,0.6292436,0.0002546354,0.00036991155,0.0010334296,0.00090597017,0.33766273,0.011447841,0.005345551,0.012068768,0.00013178904],"about_ca_topic_score_codex":0.007959174,"about_ca_topic_score_gemma":0.015109844,"teacher_disagreement_score":0.010556773,"about_ca_system_score_codex":0.0004864798,"about_ca_system_score_gemma":0.0011265315,"threshold_uncertainty_score":0.015825748},"labels":[],"label_agreement":null},{"id":"W3023437205","doi":"","title":"Investigating the relationship between evolutionary coupling and software bug-proneness.","year":2019,"lang":"en","type":"article","venue":"Conference of the Centre for Advanced Studies on Collaborative Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Software; Software bug; Programming language","score_opus":0.15121336623032364,"score_gpt":0.38980858421846754,"score_spread":0.2385952179881439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3023437205","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9982498,0.0001860925,0.00090958085,0.00007684637,0.000006961794,0.00001188324,0.00006554533,0.000017115963,0.00047611157],"genre_scores_gemma":[0.99876606,0.000045565594,0.0008408908,0.000015765168,0.0000042855336,0.000010940634,0.00009996892,0.000010315825,0.00020620263],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9949437,0.0028877927,0.00036397474,0.0005794393,0.0009409835,0.000284083],"domain_scores_gemma":[0.83808714,0.113665,0.032464437,0.004744413,0.006661089,0.004377908],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061025545,0.00033779404,0.00027005273,0.0021835712,0.00034301195,0.0010500432,0.0007449832,0.0007213971,0.0021254586],"category_scores_gemma":[0.091622815,0.00032002418,0.0003514587,0.0018316312,0.00048092508,0.001137436,0.0012026363,0.0012855412,0.00029358268],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039759427,0.00031640657,0.9790647,0.000051653973,0.00031023996,0.000105782994,0.0005423692,0.0011866762,0.0016322085,0.0002725863,0.00018821062,0.015931582],"study_design_scores_gemma":[0.000018116703,0.0005606992,0.9935846,0.000018417677,0.00012557149,0.00016708022,0.00033044044,0.004038058,0.00045276,0.00046804786,0.00022248835,0.000013785011],"about_ca_topic_score_codex":0.0030813022,"about_ca_topic_score_gemma":0.004359465,"teacher_disagreement_score":0.0061025545,"about_ca_system_score_codex":0.0004088376,"about_ca_system_score_gemma":0.0006052003,"threshold_uncertainty_score":0.03227371},"labels":[],"label_agreement":null},{"id":"W3023991692","doi":"","title":"Clones and Macro-Co-Changes","year":2014,"lang":"en","type":"article","venue":"VUBIR (Vrije Universiteit Brussel)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Macro; clone (Java method); Commit; Computer science; Cloning (programming); Biology; Genetics; Programming language; Database; Gene","score_opus":0.007667919469674493,"score_gpt":0.21427942794025745,"score_spread":0.20661150847058296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3023991692","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9554046,0.0010226846,0.035554383,0.00019883834,0.000044683686,0.00020191606,0.0008903106,0.0005997245,0.006082936],"genre_scores_gemma":[0.98689795,0.00016012852,0.010620846,0.00005783906,0.000032226628,0.00009326062,0.0005052171,0.00009850717,0.0015340046],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.98726404,0.00214688,0.001389646,0.0027722572,0.0056631276,0.0007639071],"domain_scores_gemma":[0.83842915,0.087151304,0.04133681,0.018430915,0.011740198,0.0029115935],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004334417,0.00046186443,0.0005356917,0.003886643,0.0009161962,0.0020796144,0.00093857286,0.0010301291,0.0026932599],"category_scores_gemma":[0.0579808,0.0004184944,0.0004025612,0.003751482,0.0020212922,0.0044306745,0.0022282614,0.0008910185,0.0004852927],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067306863,0.0001416946,0.85016906,0.00035510663,0.00025631802,0.0008415641,0.0032054181,0.0032747735,0.010203929,0.009627788,0.00121777,0.12003345],"study_design_scores_gemma":[0.000034145673,0.00066939753,0.92118275,0.00016061694,0.00026242348,0.004594386,0.003309984,0.01989269,0.01673312,0.017881928,0.015165828,0.00011279388],"about_ca_topic_score_codex":0.0027953736,"about_ca_topic_score_gemma":0.0028651895,"teacher_disagreement_score":0.004334417,"about_ca_system_score_codex":0.0009731663,"about_ca_system_score_gemma":0.00086273317,"threshold_uncertainty_score":0.022922873},"labels":[],"label_agreement":null},{"id":"W3024356476","doi":"10.1007/s10664-020-09878-9","title":"On the time-based conclusion stability of cross-project defect prediction models","year":2020,"lang":"en","type":"preprint","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Stability (learning theory); Computer science; Software; Predictive modelling; Product (mathematics); Limit (mathematics); Data mining; Time limit; Analytics; Data science; Empirical research; Econometrics; Reliability engineering; Statistics; Machine learning; Mathematics; Engineering; Systems engineering","score_opus":0.07151071104906312,"score_gpt":0.31989840498568256,"score_spread":0.24838769393661944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3024356476","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37764737,0.0025657122,0.6092844,0.0035184233,0.0002443677,0.000094051364,0.0005944589,0.0006364219,0.0054148664],"genre_scores_gemma":[0.9808966,0.00041313548,0.015994022,0.00022463051,0.0001351492,0.00004480467,0.0005545118,0.00016740631,0.0015696423],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9942,0.003392561,0.00023518606,0.0011983004,0.0006909049,0.0002830732],"domain_scores_gemma":[0.7162458,0.2610842,0.0064616073,0.006342258,0.008508021,0.0013581788],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.027729636,0.0009094398,0.0016086847,0.0022694631,0.000936244,0.0028155996,0.0024940763,0.002226517,0.003733435],"category_scores_gemma":[0.16730647,0.0006561574,0.0012672795,0.0013499568,0.0020213008,0.004309156,0.0025868302,0.0039507058,0.00047189195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012034791,0.0002073446,0.02870645,0.0002430945,0.00053214683,0.00026716513,0.0004770511,0.83459824,0.002213557,0.06482403,0.003430249,0.06329714],"study_design_scores_gemma":[0.000009452336,0.00003723403,0.0013962862,0.0000232581,0.000030152083,0.000017777978,0.000028633829,0.98668706,0.00032832622,0.011322481,0.00010971967,0.00000970578],"about_ca_topic_score_codex":0.0069598425,"about_ca_topic_score_gemma":0.0035221933,"teacher_disagreement_score":0.97227037,"about_ca_system_score_codex":0.0017083236,"about_ca_system_score_gemma":0.001323066,"threshold_uncertainty_score":0.14665008},"labels":[],"label_agreement":null},{"id":"W3024531409","doi":"10.1109/iraset48871.2020.9092144","title":"A Sketch of a Deep Learning Approach for Discovering UML Class Diagrams from System’s Textual Specification","year":2020,"lang":"en","type":"article","venue":"2020 1st International Conference on Innovative Research in Applied Science, Engineering and Technology (IRASET)","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Rimouski","funders":"","keywords":"Computer science; Sketch; Artificial intelligence; Unified Modeling Language; Formal specification; Class diagram; Natural language processing; Programming language; Context (archaeology); Deep learning; Class (philosophy); Machine learning; Software","score_opus":0.057184414616548315,"score_gpt":0.31856250425302046,"score_spread":0.26137808963647213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3024531409","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008420431,0.0004383396,0.99456906,0.0008454718,0.000038017868,0.0001109088,0.00026954297,0.0012001556,0.0016865734],"genre_scores_gemma":[0.03931532,0.0017240848,0.9504066,0.00073802995,0.00006206203,0.00041405234,0.0011269626,0.00015727896,0.006055661],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99942195,0.00014799176,0.00007807272,0.00016734387,0.00013971653,0.00004495687],"domain_scores_gemma":[0.99929523,0.00032504034,0.000045004097,0.00014267466,0.00014644807,0.00004562453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013182058,0.0011745453,0.00058459054,0.0016848223,0.00055804546,0.0027137967,0.0025622305,0.0022985777,0.008299012],"category_scores_gemma":[0.0028473565,0.00085449213,0.0013558258,0.0011261012,0.0011017176,0.0030848095,0.0017425681,0.0030684974,0.0031781767],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018402064,0.0002580226,0.0016690051,0.0012362209,0.00018100515,0.0006204309,0.00057935016,0.11778374,0.01778769,0.19502476,0.01807937,0.64659643],"study_design_scores_gemma":[0.00004160622,0.0001583745,0.0007181096,0.00039984466,0.000071360475,0.0007070707,0.00012669683,0.70820695,0.010530682,0.205977,0.07298953,0.00007277009],"about_ca_topic_score_codex":0.0054906984,"about_ca_topic_score_gemma":0.007872777,"teacher_disagreement_score":0.008299012,"about_ca_system_score_codex":0.0012263218,"about_ca_system_score_gemma":0.0012847376,"threshold_uncertainty_score":0.02776301},"labels":[],"label_agreement":null},{"id":"W3025537827","doi":"10.1109/tse.2021.3087087","title":"Generating Unit Tests for Documentation","year":2021,"lang":"en","type":"preprint","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Internal documentation; Unit testing; Computer science; Software documentation; Redundancy (engineering); Artifact (error); Source code; Software engineering; Software; Database; Programming language; Operating system; Software development; Artificial intelligence; Software development process; Software construction","score_opus":0.03303154572484931,"score_gpt":0.29791213454178017,"score_spread":0.2648805888169309,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3025537827","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.096995,0.00033441148,0.8557764,0.00028487298,0.00020577968,0.0007329653,0.001840688,0.03698016,0.0068497336],"genre_scores_gemma":[0.30035293,0.00019444784,0.6826017,0.0002303398,0.00007078141,0.0007251151,0.0056124646,0.006055742,0.004156501],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99159086,0.0028214757,0.000867408,0.0011295043,0.0032139353,0.00037680118],"domain_scores_gemma":[0.92666346,0.04100972,0.0051042875,0.0151686575,0.0112155685,0.00083830376],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004975728,0.0014381963,0.0010553538,0.0042161867,0.00050228596,0.0020287924,0.0020131979,0.0014484953,0.0059088278],"category_scores_gemma":[0.05983199,0.00093481835,0.0012752995,0.002070534,0.00077377044,0.001818247,0.0020320544,0.0011553483,0.002950515],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007447159,0.00064666447,0.022070628,0.0012290626,0.00018790046,0.0017485793,0.0012860312,0.04636938,0.05480761,0.023118045,0.025785325,0.82200617],"study_design_scores_gemma":[0.00043511565,0.0010599777,0.007760834,0.0005326889,0.0002315466,0.0027963032,0.00041267605,0.54669374,0.3340314,0.03947902,0.06634304,0.00022368548],"about_ca_topic_score_codex":0.00092455465,"about_ca_topic_score_gemma":0.0010407611,"teacher_disagreement_score":0.0059088278,"about_ca_system_score_codex":0.0007804861,"about_ca_system_score_gemma":0.0017661856,"threshold_uncertainty_score":0.026314437},"labels":[],"label_agreement":null},{"id":"W3028200206","doi":"10.1007/s10664-020-09837-4","title":"Do code review measures explain the incidence of post-release defects?","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Operationalization; Code (set theory); Computer science; Replication (statistics); Code review; Replicate; Variable (mathematics); Empirical research; Contrast (vision); Econometrics; Statistics; Software; Software quality; Artificial intelligence; Mathematics; Software development; Programming language","score_opus":0.042557336039330304,"score_gpt":0.29470602258305645,"score_spread":0.25214868654372613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3028200206","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99370366,0.0012042754,0.0015412857,0.0007922967,0.000048163292,0.00002815654,0.0008260564,0.000067120476,0.0017888768],"genre_scores_gemma":[0.99890053,0.000112194626,0.00022055203,0.000068714326,0.000023710196,0.000009732918,0.00029750384,0.00002789455,0.00033910535],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99234444,0.0022400753,0.0011326843,0.001089467,0.002406876,0.00078647124],"domain_scores_gemma":[0.5384732,0.24040405,0.18205297,0.017546702,0.01634718,0.005175964],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011275735,0.0005604706,0.0006285384,0.0049251085,0.00042395227,0.0021399518,0.0016613622,0.0018356931,0.004334835],"category_scores_gemma":[0.18100414,0.0005366016,0.0010331775,0.003806828,0.0012024358,0.0030761212,0.00095748954,0.0017888322,0.0010394948],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010748851,0.00008627708,0.99232,0.000054815482,0.0001967927,0.00004091048,0.00017511536,0.00026760763,0.00020183757,0.0001790557,0.00035618414,0.0060139503],"study_design_scores_gemma":[0.000008596455,0.00010447547,0.9978309,0.00003765348,0.00006558553,0.000080918,0.00018170472,0.00086086255,0.00023345857,0.00026948986,0.000315886,0.000010359834],"about_ca_topic_score_codex":0.005785036,"about_ca_topic_score_gemma":0.009438024,"teacher_disagreement_score":0.9887243,"about_ca_system_score_codex":0.0008661799,"about_ca_system_score_gemma":0.0013501208,"threshold_uncertainty_score":0.05963254},"labels":[],"label_agreement":null},{"id":"W302904365","doi":"10.1007/978-3-319-14130-5_23","title":"Mining Software Components from Object-Oriented APIs","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Software engineering; Programming language","score_opus":0.018495793163478054,"score_gpt":0.24491322568238794,"score_spread":0.2264174325189099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W302904365","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20566949,0.0045292084,0.75253826,0.00060138723,0.0002568361,0.00073303445,0.006629565,0.019733943,0.009308336],"genre_scores_gemma":[0.2786863,0.0027649333,0.68039584,0.00013479816,0.000130937,0.00037868856,0.028448155,0.0017401401,0.0073201214],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990946,0.00007707265,0.00007746452,0.00020833104,0.00046041165,0.000082181614],"domain_scores_gemma":[0.99825865,0.0007684823,0.0002054981,0.00032928065,0.00037251803,0.000065623404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00052683975,0.0019106386,0.0009748172,0.006478539,0.0006762646,0.0021918276,0.0020710356,0.00090788695,0.0018052326],"category_scores_gemma":[0.003907998,0.00076506444,0.0027985906,0.005711096,0.00042156182,0.0025011972,0.0010433497,0.001228646,0.0020030807],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025900156,0.00036316432,0.02603188,0.001224854,0.0002833837,0.0010624705,0.00042469345,0.018329073,0.021242633,0.009223701,0.014208584,0.9073466],"study_design_scores_gemma":[0.00012514101,0.00033946495,0.03184331,0.0005913044,0.0010301848,0.0025975187,0.0012590361,0.7673787,0.055643614,0.075381,0.06369295,0.00011778466],"about_ca_topic_score_codex":0.0043506073,"about_ca_topic_score_gemma":0.0077595413,"teacher_disagreement_score":0.006478539,"about_ca_system_score_codex":0.0004449575,"about_ca_system_score_gemma":0.0014182761,"threshold_uncertainty_score":0.008650541},"labels":[],"label_agreement":null},{"id":"W3029302266","doi":"10.1109/vl/hcc50065.2020.9127202","title":"Code Duplication and Reuse in Jupyter Notebooks","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Reuse; Computer science; Code (set theory); Code reuse; Software; Software engineering; Source code; Sample (material); Programming language; Engineering","score_opus":0.04363914024277774,"score_gpt":0.29553567062364006,"score_spread":0.2518965303808623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3029302266","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97898304,0.0004929383,0.015061769,0.00036346578,0.000022537035,0.000094856056,0.00023159101,0.0016904889,0.0030594002],"genre_scores_gemma":[0.9639039,0.00033534155,0.029208293,0.0001916054,0.000027430178,0.00013687531,0.00069127925,0.0009849259,0.0045202514],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9911179,0.0025643543,0.00063798175,0.001402426,0.0037301711,0.00054710923],"domain_scores_gemma":[0.90409505,0.055397384,0.016758341,0.015288261,0.00620493,0.0022560444],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.006941642,0.000612653,0.0005753954,0.003847,0.0017209878,0.0035084244,0.0018553052,0.0010009463,0.002514607],"category_scores_gemma":[0.08158562,0.00095383474,0.00048659643,0.0035514103,0.0026163543,0.0071861153,0.0050201663,0.0011728342,0.0005621556],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018942914,0.000618783,0.40471235,0.0016621663,0.0002638693,0.007541803,0.21078372,0.0035648681,0.030198827,0.009450188,0.011653924,0.31765515],"study_design_scores_gemma":[0.00017243781,0.0016782493,0.6487384,0.0017684597,0.00038687408,0.017078808,0.086002156,0.024979243,0.05262815,0.017136158,0.14877127,0.000659786],"about_ca_topic_score_codex":0.003191418,"about_ca_topic_score_gemma":0.005491337,"teacher_disagreement_score":0.9930584,"about_ca_system_score_codex":0.001594993,"about_ca_system_score_gemma":0.0017496394,"threshold_uncertainty_score":0.036711395},"labels":[],"label_agreement":null},{"id":"W3032170634","doi":"10.1109/tse.2020.2998503","title":"Automatic Generation of Acceptance Test Cases From Use Case Specifications: An NLP-Based Approach","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":85,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"European Research Council; European Commission","keywords":"Computer science; Executable; Test case; Software requirements specification; Acceptance testing; Software engineering; Test script; System under test; Conformance testing; Test Management Approach; Test (biology); Formal specification; Reliability engineering; Software; Software system; Programming language; Machine learning; Software construction","score_opus":0.11804880023702771,"score_gpt":0.2662858943771059,"score_spread":0.1482370941400782,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3032170634","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018216228,0.00012319794,0.96683425,0.00018535838,0.000033443557,0.0007333611,0.00060079165,0.010876961,0.0023962602],"genre_scores_gemma":[0.15627919,0.00021447147,0.8348013,0.00019829486,0.000027954362,0.0011143199,0.003714979,0.0018287658,0.0018207087],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9922091,0.003222575,0.00055735477,0.0010467963,0.0026175063,0.00034670354],"domain_scores_gemma":[0.97451514,0.018245103,0.0019695573,0.002090059,0.0029402403,0.00023982154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040108664,0.002351195,0.0009876164,0.003688694,0.0007067601,0.0022708867,0.0023693317,0.00205882,0.004802352],"category_scores_gemma":[0.024978606,0.0013680565,0.002370307,0.0014788869,0.0013522003,0.0015896691,0.0020797905,0.0017525019,0.0021299284],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070131524,0.0012489951,0.009812379,0.0020404777,0.0003739326,0.005674814,0.0036217067,0.24068002,0.10256836,0.03426008,0.013172232,0.58584577],"study_design_scores_gemma":[0.00014874198,0.00021771003,0.0013108127,0.00020354285,0.00014051993,0.001003115,0.0003551022,0.91228515,0.056300826,0.012903058,0.015046919,0.00008457332],"about_ca_topic_score_codex":0.0037135738,"about_ca_topic_score_gemma":0.0038034425,"teacher_disagreement_score":0.004802352,"about_ca_system_score_codex":0.0011407291,"about_ca_system_score_gemma":0.0019599276,"threshold_uncertainty_score":0.021211743},"labels":[],"label_agreement":null},{"id":"W3033176747","doi":"10.48550/arxiv.2006.03103","title":"SMIE: Weakness is Power!: Auto-indentation with incomplete information","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Grammar; Parsing; Computer science; Code (set theory); Indentation; Simple (philosophy); Programming language; Face (sociological concept); Power (physics); Artificial intelligence; Natural language processing; Linguistics; Physics; Set (abstract data type)","score_opus":0.05255077302078416,"score_gpt":0.19103051793586023,"score_spread":0.1384797449150761,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3033176747","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009161278,0.00038843742,0.9126219,0.0018698346,0.00042355276,0.00009351411,0.0003805206,0.06447866,0.010582267],"genre_scores_gemma":[0.17226374,0.0006680886,0.76606214,0.0028044148,0.00046846474,0.0002462702,0.0019697868,0.03147318,0.02404391],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9933624,0.001757512,0.000448241,0.0011801685,0.0028134596,0.00043818622],"domain_scores_gemma":[0.98226273,0.0063653896,0.00087071914,0.008388903,0.0017388054,0.0003734301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0081026545,0.0015448484,0.0012110601,0.0017785855,0.0014656397,0.0045792446,0.004973816,0.0025139085,0.011923018],"category_scores_gemma":[0.035511866,0.0016403486,0.0015665017,0.0015771489,0.0043618996,0.01328903,0.010769942,0.006010272,0.007904116],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079930195,0.0002099252,0.004585608,0.00093107205,0.00024337576,0.00096737314,0.0030836328,0.011711145,0.030828815,0.30590078,0.1228263,0.5179127],"study_design_scores_gemma":[0.000091630696,0.00018594126,0.001220272,0.0004064124,0.00014950421,0.0016951594,0.0004114,0.18157601,0.13406412,0.34536132,0.33457825,0.00025998364],"about_ca_topic_score_codex":0.00071042863,"about_ca_topic_score_gemma":0.0012082873,"teacher_disagreement_score":0.011923018,"about_ca_system_score_codex":0.00093010516,"about_ca_system_score_gemma":0.0015648111,"threshold_uncertainty_score":0.042851448},"labels":[],"label_agreement":null},{"id":"W3037099619","doi":"10.1109/tse.2020.3004525","title":"Why Do Software Developers Use Static Analysis Tools? A User-Centered Study of Developer Needs and Motivations","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":76,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Heinz Nixdorf Stiftung; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Usability; Software engineering; Software development; Software; Static program analysis; Secure coding; World Wide Web; Human–computer interaction; Software security assurance; Computer security; Programming language; Information security","score_opus":0.04324661639113619,"score_gpt":0.24984638888479016,"score_spread":0.20659977249365397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3037099619","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91442186,0.0020586904,0.039648477,0.018315265,0.00011475785,0.00015350689,0.000103665974,0.00055006205,0.024633806],"genre_scores_gemma":[0.9831381,0.00071159005,0.011073067,0.0019499101,0.00004112423,0.000118888696,0.00007261216,0.00020689382,0.0026878868],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.96248454,0.021245819,0.0017916844,0.0022950177,0.009986603,0.0021962563],"domain_scores_gemma":[0.7920291,0.14190274,0.015734402,0.0075311484,0.037223857,0.0055786474],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.02803061,0.0008497949,0.00057800306,0.0052990187,0.003396101,0.006828937,0.0017762918,0.0035153849,0.0013875637],"category_scores_gemma":[0.12884767,0.0016153859,0.00047759473,0.002468004,0.004267436,0.011321725,0.0034341624,0.0023317325,0.0009285348],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026645794,0.00043001122,0.30056092,0.0005924513,0.00011180793,0.0017116584,0.4993075,0.0004421394,0.0086118905,0.02619916,0.00800587,0.15376006],"study_design_scores_gemma":[0.00014922905,0.0006244833,0.22668,0.002227695,0.00027631884,0.0059570377,0.5423213,0.014821085,0.010991545,0.05685009,0.13839242,0.00070875575],"about_ca_topic_score_codex":0.004187184,"about_ca_topic_score_gemma":0.0066067907,"teacher_disagreement_score":0.97196937,"about_ca_system_score_codex":0.0022186844,"about_ca_system_score_gemma":0.0037094322,"threshold_uncertainty_score":0.14824176},"labels":[],"label_agreement":null},{"id":"W3048030566","doi":"10.1145/3385732","title":"Automatic Detection of Usability Problem Encounters in Think-aloud Sessions","year":2020,"lang":"en","type":"article","venue":"ACM Transactions on Interactive Intelligent Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Usability; Think aloud protocol; Computer science; Human–computer interaction; Workflow; Usability lab; Usability inspection; Cognitive walkthrough; Pluralistic walkthrough; Usability engineering; Test (biology); Automation; Usability goals; Engineering; Database","score_opus":0.030562041448167127,"score_gpt":0.2917840176591617,"score_spread":0.2612219762109945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3048030566","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6922636,0.00031436587,0.29048643,0.00037893662,0.00015820685,0.0043870704,0.00090569456,0.0069926204,0.004113098],"genre_scores_gemma":[0.72490585,0.00023121358,0.26556656,0.00029381146,0.00007087416,0.004839069,0.0010013573,0.0006909299,0.0024003796],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9779584,0.01273931,0.0019511079,0.0028712475,0.0038836582,0.00059627736],"domain_scores_gemma":[0.76383275,0.18358617,0.016862303,0.0122774765,0.021462409,0.0019789066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016310632,0.00209341,0.0012536927,0.0027904012,0.00094542984,0.0028123038,0.0016834114,0.0012629951,0.0022670147],"category_scores_gemma":[0.10897759,0.000899051,0.00057269854,0.0011039354,0.0009464614,0.0026577068,0.0019622834,0.0017929417,0.0019060767],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043253032,0.0029916214,0.0775196,0.004317234,0.00036278766,0.00075713044,0.103225544,0.0032232834,0.25782445,0.0025337012,0.007648355,0.5352709],"study_design_scores_gemma":[0.0011463278,0.01108131,0.2965123,0.0019816707,0.0006968067,0.002541598,0.053082675,0.1937846,0.37932813,0.02240613,0.03582353,0.0016149633],"about_ca_topic_score_codex":0.0006625039,"about_ca_topic_score_gemma":0.0014439368,"teacher_disagreement_score":0.016310632,"about_ca_system_score_codex":0.000797352,"about_ca_system_score_gemma":0.0013090221,"threshold_uncertainty_score":0.0862599},"labels":[],"label_agreement":null},{"id":"W3048169489","doi":"10.1109/icsme46990.2020.00097","title":"DR-Tools: a suite of lightweight open-source tools to measure and visualize Java source code","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Suite; Computer science; Java; Source code; Software engineering; Heuristics; Open source; Set (abstract data type); Code review; Code (set theory); Codebase; Software suite; Software evolution; Software; Software quality; Programming language; Software development; Operating system; Software construction","score_opus":0.08766148667862791,"score_gpt":0.3302400270729327,"score_spread":0.2425785403943048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3048169489","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017623395,0.0011734434,0.57973975,0.0003005022,0.00016559294,0.0007300808,0.009875572,0.38264352,0.0077481545],"genre_scores_gemma":[0.14236638,0.0015149865,0.7415356,0.00041703775,0.0001573881,0.0016865274,0.036604747,0.06625103,0.009466315],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946426,0.0010194394,0.00060025277,0.0007004797,0.0027446803,0.0002925924],"domain_scores_gemma":[0.98342,0.007910811,0.0026267795,0.0028807765,0.0024905282,0.00067108637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033122213,0.002574633,0.00084151054,0.009914826,0.0006417775,0.0025351776,0.0027007808,0.0012835059,0.007868182],"category_scores_gemma":[0.029995278,0.001471596,0.0013956538,0.004121752,0.00070634735,0.004266494,0.0045525315,0.002273865,0.0058670435],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006567209,0.0006276756,0.017093474,0.0024239437,0.00039470143,0.0009679463,0.0024358036,0.01433479,0.034476005,0.017395062,0.15688933,0.75230455],"study_design_scores_gemma":[0.00071405334,0.0009644171,0.06156202,0.0018590301,0.0003671312,0.0033436893,0.0010802504,0.24722648,0.11688712,0.07972168,0.48501602,0.0012580865],"about_ca_topic_score_codex":0.0032823244,"about_ca_topic_score_gemma":0.004929953,"teacher_disagreement_score":0.009914826,"about_ca_system_score_codex":0.00063359115,"about_ca_system_score_gemma":0.0014203496,"threshold_uncertainty_score":0.02632165},"labels":[],"label_agreement":null},{"id":"W3048218005","doi":"10.1145/3379597.3387508","title":"Multi-language Design Smells","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code smell; Computer science; Relevance (law); Software engineering; Context (archaeology); Quality (philosophy); Software quality; Software development; Plan (archaeology); Software; Software system; Software design; Data science; World Wide Web; Programming language","score_opus":0.06335920762229547,"score_gpt":0.2918802640990534,"score_spread":0.2285210564767579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3048218005","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9722493,0.0009033474,0.022537148,0.0004906132,0.000030743628,0.00011590236,0.0005189145,0.00066513853,0.0024889102],"genre_scores_gemma":[0.9822327,0.00036391732,0.0142907305,0.00020579397,0.000020928817,0.000097832286,0.00084037596,0.00023297973,0.0017146987],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99032545,0.0025429728,0.0015302217,0.0015514669,0.0034909034,0.0005589794],"domain_scores_gemma":[0.86717826,0.066589676,0.040654536,0.012809928,0.01054846,0.0022191312],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008460927,0.000671843,0.00056004606,0.005355066,0.00097510044,0.0019251133,0.000824065,0.0008413173,0.0013346215],"category_scores_gemma":[0.04421101,0.0005338831,0.0006993418,0.004007417,0.0013026586,0.003376529,0.0030672858,0.0010696867,0.0004030214],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045302682,0.000265265,0.75363916,0.0022725796,0.00022346863,0.0055114515,0.043532636,0.0018696707,0.027521469,0.0028823544,0.0027882028,0.15904076],"study_design_scores_gemma":[0.0000457822,0.00055552775,0.8999062,0.0013735355,0.0002859061,0.008725147,0.01726068,0.008605491,0.01812735,0.0063331216,0.03856835,0.0002128523],"about_ca_topic_score_codex":0.0021489556,"about_ca_topic_score_gemma":0.004110751,"teacher_disagreement_score":0.008460927,"about_ca_system_score_codex":0.0012662828,"about_ca_system_score_gemma":0.001143943,"threshold_uncertainty_score":0.04474616},"labels":[],"label_agreement":null},{"id":"W3048392083","doi":"10.1186/s13173-020-00100-8","title":"MylynSDP — Process - aware artifact filtering based on interest","year":2020,"lang":"en","type":"article","venue":"Journal of the Brazilian Computer Society","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Coordenação de Aperfeiçoamento de Pessoal de Nível Superior","keywords":"Computer science; Artifact (error); Software engineering; Software development; Software; Function (biology); Plug-in; Process (computing); Task (project management); Context (archaeology); Software development process; Goal-Driven Software Development Process; Programming language; Artificial intelligence; Systems engineering; Engineering","score_opus":0.03342834874107232,"score_gpt":0.2720713729770696,"score_spread":0.23864302423599731,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3048392083","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14705797,0.00059346296,0.8233276,0.00029551628,0.00006823411,0.0005725093,0.0008741749,0.023974186,0.003236336],"genre_scores_gemma":[0.6257312,0.00024299209,0.36695588,0.00013562481,0.00004743186,0.00039155182,0.0021577843,0.0006697892,0.0036677716],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99783844,0.00039393787,0.00017207589,0.0006191165,0.0008453804,0.00013101783],"domain_scores_gemma":[0.9939255,0.0026572503,0.00085987174,0.0009181447,0.0011991424,0.00044010484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002226835,0.0010499456,0.0007730644,0.0032403173,0.0004863175,0.0017133736,0.0013590292,0.0006282542,0.0018741051],"category_scores_gemma":[0.007980159,0.00049184,0.00092958286,0.0014358398,0.00036973305,0.0024670654,0.0025604104,0.0008174248,0.00094717444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013249328,0.0010609879,0.04775524,0.0009366445,0.00015002427,0.0006592409,0.0017407361,0.0059563387,0.062285688,0.0041375556,0.0082313055,0.8657613],"study_design_scores_gemma":[0.00037000206,0.0022102536,0.13651559,0.00032575088,0.0005351975,0.0020444323,0.0014080339,0.62867707,0.15637472,0.0149144055,0.056300763,0.0003238756],"about_ca_topic_score_codex":0.0016600601,"about_ca_topic_score_gemma":0.0019786714,"teacher_disagreement_score":0.0032403173,"about_ca_system_score_codex":0.0005429692,"about_ca_system_score_gemma":0.00076059165,"threshold_uncertainty_score":0.011776745},"labels":[],"label_agreement":null},{"id":"W3059397974","doi":"10.1016/j.infsof.2020.106392","title":"Predicting continuous integration build failures using evolutionary search","year":2020,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Computer science; Benchmark (surveying); Genetic programming; Context (archaeology); Process (computing); Software; Software engineering; Resource (disambiguation); Search-based software engineering; Machine learning; Outcome (game theory); Software development; Artificial intelligence; Data mining; Software development process; Programming language","score_opus":0.014918686799295273,"score_gpt":0.24796759159791096,"score_spread":0.23304890479861567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3059397974","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94197905,0.00031058205,0.05553288,0.0001230461,0.000022892138,0.00003552305,0.00019281461,0.00044223387,0.0013609926],"genre_scores_gemma":[0.9888148,0.000040463583,0.010348053,0.000010720013,0.0000041298676,0.000013412699,0.0002133126,0.00002344449,0.00053172756],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99948394,0.00012358949,0.00003082885,0.000118647826,0.00016383834,0.00007905964],"domain_scores_gemma":[0.99460006,0.0037850775,0.0005227219,0.00025122048,0.00063187553,0.00020900452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012613209,0.00094707205,0.000709823,0.0025757807,0.00039601265,0.00080934126,0.0010907914,0.0012868238,0.0014993687],"category_scores_gemma":[0.0076184045,0.0004654153,0.00055866624,0.0014103007,0.00038816786,0.001231017,0.0006246896,0.0009346077,0.00031867513],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027823242,0.00028082382,0.05223165,0.000047781614,0.00011120805,0.00014957544,0.00006415392,0.8889114,0.0017539192,0.0006551252,0.0005871085,0.054929025],"study_design_scores_gemma":[0.000004347995,0.000039095587,0.0023748705,0.0000033370186,0.000010700112,0.000016224796,0.000014330217,0.99696714,0.00020990029,0.00032395974,0.000033371824,0.0000027168765],"about_ca_topic_score_codex":0.009948942,"about_ca_topic_score_gemma":0.009976041,"teacher_disagreement_score":0.009948942,"about_ca_system_score_codex":0.0006704329,"about_ca_system_score_gemma":0.00062542385,"threshold_uncertainty_score":0.019782066},"labels":[],"label_agreement":null},{"id":"W3074372486","doi":"10.11606/t.55.2020.tde-18082020-163540","title":"Classificação automática de questões baseada em competências: ENEM - Estudo de caso","year":2020,"lang":"pt","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Personalization; Computer science; The Internet; Point (geometry); Artificial intelligence; Question answering; Information retrieval; Subject (documents); Natural language processing; World Wide Web","score_opus":0.028425327744483038,"score_gpt":0.30287238698437463,"score_spread":0.2744470592398916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3074372486","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73024905,0.0058455584,0.2064069,0.001382552,0.0006292393,0.0012222355,0.006670828,0.007379956,0.040213674],"genre_scores_gemma":[0.8840144,0.0009528524,0.10326899,0.00015295605,0.00016694583,0.0002777082,0.0044215554,0.000334928,0.0064097135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99589455,0.0010182457,0.00036122053,0.0012324018,0.001199304,0.00029419007],"domain_scores_gemma":[0.98954445,0.005821748,0.0005237575,0.0011081387,0.0027141324,0.00028782734],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045529096,0.0010445516,0.00094123575,0.0063630426,0.0009007436,0.002990792,0.0013866362,0.0011379176,0.003813657],"category_scores_gemma":[0.022569006,0.00040170155,0.0011172806,0.0029478706,0.0008593799,0.0023963968,0.0016145871,0.0009610686,0.0019273224],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011526608,0.00092572015,0.12595417,0.0012364088,0.00028852216,0.0016879812,0.0044228504,0.005966411,0.013548443,0.0043535368,0.018460408,0.8220028],"study_design_scores_gemma":[0.00035776754,0.0017759023,0.22404034,0.0014593791,0.0012708905,0.008440952,0.012931079,0.530654,0.05739228,0.018947806,0.1423068,0.00042282874],"about_ca_topic_score_codex":0.012757158,"about_ca_topic_score_gemma":0.01674063,"teacher_disagreement_score":0.012757158,"about_ca_system_score_codex":0.00089559046,"about_ca_system_score_gemma":0.0011966919,"threshold_uncertainty_score":0.02536577},"labels":[],"label_agreement":null},{"id":"W3081422979","doi":"10.1007/s10270-020-00823-4","title":"Consistent change propagation within models","year":2020,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Ministère de l'Agriculture, des Pêcheries et de l'Alimentation","funders":"Natural Sciences and Engineering Research Council of Canada; Österreichische Forschungsförderungsgesellschaft; Austrian Science Fund","keywords":"Computer science; Code refactoring; Focus (optics); Data science; Risk analysis (engineering); Software engineering; Software; Programming language","score_opus":0.1313204512815091,"score_gpt":0.26692143037497357,"score_spread":0.13560097909346447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3081422979","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0123795215,0.00049378164,0.9750114,0.0010829921,0.00009755842,0.00059866393,0.00045990146,0.0054404195,0.004435734],"genre_scores_gemma":[0.12304649,0.000514933,0.8687774,0.00055004796,0.00007529633,0.0005013655,0.0014738328,0.001898537,0.0031620413],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9673963,0.00911813,0.0033119055,0.007056284,0.011706455,0.0014108948],"domain_scores_gemma":[0.92821515,0.030921398,0.005230739,0.025873367,0.008582289,0.0011770795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019171242,0.0029575725,0.0018877983,0.0061959373,0.002870053,0.010990281,0.00852768,0.0060500847,0.005869544],"category_scores_gemma":[0.08364007,0.0040785875,0.006248933,0.0038957072,0.0054097497,0.022835193,0.013950115,0.0067812097,0.0021165453],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005202197,0.0006386271,0.013828589,0.001937237,0.00069374766,0.0025133868,0.008238873,0.18641701,0.015971275,0.3114781,0.011814186,0.4459487],"study_design_scores_gemma":[0.00017271317,0.00037015125,0.0015576849,0.000959498,0.0008131563,0.0015546684,0.0018492513,0.44914046,0.025565015,0.389238,0.12842943,0.00035000226],"about_ca_topic_score_codex":0.0068517583,"about_ca_topic_score_gemma":0.0072357464,"teacher_disagreement_score":0.019171242,"about_ca_system_score_codex":0.0038421294,"about_ca_system_score_gemma":0.007144656,"threshold_uncertainty_score":0.101388395},"labels":[],"label_agreement":null},{"id":"W3081943439","doi":"10.1007/s10664-020-09863-2","title":"CROKAGE: effective solution recommendation for programming tasks by leveraging crowd knowledge","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Leverage (statistics); Code (set theory); Information retrieval; Task (project management); Relevance (law); Programming language; Artificial intelligence","score_opus":0.03546418867217926,"score_gpt":0.3080167954479544,"score_spread":0.27255260677577514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3081943439","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16161682,0.006738643,0.75025326,0.0019930042,0.0012420157,0.001461876,0.0048623495,0.052109867,0.019722156],"genre_scores_gemma":[0.4182077,0.00083497394,0.5584208,0.00078086334,0.00028258134,0.0006767238,0.0077180834,0.0010032548,0.012075054],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99830544,0.00039678058,0.00006344386,0.00054626475,0.0005312573,0.00015676426],"domain_scores_gemma":[0.99722666,0.0013028004,0.00014748119,0.0006167785,0.00048095433,0.0002252318],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017689326,0.0020507914,0.0016616777,0.0056689293,0.001287634,0.0015628159,0.0023172451,0.0028520324,0.00532426],"category_scores_gemma":[0.009762625,0.0007053831,0.001127334,0.002604997,0.0006490595,0.0034459927,0.0025669083,0.0018348554,0.0026239667],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016225956,0.002131168,0.010572235,0.0007313925,0.000523725,0.00027711567,0.0003885941,0.06307398,0.014436991,0.0050375033,0.091353066,0.8098516],"study_design_scores_gemma":[0.0002480451,0.0002829508,0.0018632733,0.00006461775,0.000113322705,0.00011702862,0.00014487894,0.9708827,0.005056873,0.009590792,0.011568202,0.00006726633],"about_ca_topic_score_codex":0.016550018,"about_ca_topic_score_gemma":0.039880674,"teacher_disagreement_score":0.016550018,"about_ca_system_score_codex":0.0008308554,"about_ca_system_score_gemma":0.0020176799,"threshold_uncertainty_score":0.032907367},"labels":[],"label_agreement":null},{"id":"W3082966820","doi":"10.1007/s10664-021-10005-5","title":"Understanding peer review of software engineering papers","year":2021,"lang":"en","type":"preprint","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Novelty; Quality (philosophy); Computer science; Technical peer review; Peer review; Software technical review; Psychology; Medical education; Software; Engineering ethics; Software quality; Software development; Engineering; Medicine; Political science; Social psychology","score_opus":0.10291582984899653,"score_gpt":0.3172505489139484,"score_spread":0.21433471906495186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3082966820","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16772424,0.06764422,0.20906317,0.18802309,0.014767412,0.0011476023,0.0023949435,0.0028438638,0.34639147],"genre_scores_gemma":[0.90028065,0.015762819,0.03147718,0.005168047,0.008312127,0.00035692626,0.0021431227,0.0011278685,0.035371337],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.898334,0.057304997,0.005584836,0.0056114956,0.030560583,0.0026040229],"domain_scores_gemma":[0.33479303,0.47737786,0.039020002,0.030546512,0.110902734,0.0073598633],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06389796,0.0007631439,0.001471231,0.010776294,0.003957869,0.02156978,0.0027113883,0.0058126887,0.023699166],"category_scores_gemma":[0.5207444,0.00084052805,0.0009225585,0.008372445,0.0038793301,0.025347808,0.0054076,0.0035289226,0.0054954733],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042761868,0.00026156858,0.029271258,0.0028173807,0.0005214743,0.0010794471,0.022766251,0.0047422787,0.0020155269,0.38106373,0.20886226,0.3461712],"study_design_scores_gemma":[0.00019267293,0.00018990683,0.021116583,0.0017135747,0.00032327414,0.000727453,0.010026106,0.015496283,0.0024321128,0.5948268,0.35273582,0.00021953175],"about_ca_topic_score_codex":0.003804906,"about_ca_topic_score_gemma":0.0031999422,"teacher_disagreement_score":0.93610203,"about_ca_system_score_codex":0.004606027,"about_ca_system_score_gemma":0.010072398,"threshold_uncertainty_score":0.3379287},"labels":[],"label_agreement":null},{"id":"W3084301791","doi":"10.1007/s10664-020-09864-1","title":"Automated demarcation of requirements in textual specifications: a machine learning-based approach","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"H2020 European Research Council; Natural Sciences and Engineering Research Council of Canada; Fonds National de la Recherche Luxembourg; European Commission","keywords":"Computer science; Software requirements specification; Requirements analysis; Formal specification; Markup language; System requirements specification; Simple (philosophy); Identifier; Task (project management); Requirements engineering; Variety (cybernetics); Software engineering; Enforcement; Artificial intelligence; Programming language; XML; Systems engineering; Engineering; Software; World Wide Web","score_opus":0.07854626844973958,"score_gpt":0.29834454275323014,"score_spread":0.21979827430349055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3084301791","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10854295,0.0023529162,0.8090057,0.0017667296,0.00018626933,0.0009140429,0.008728311,0.062326413,0.006176751],"genre_scores_gemma":[0.19577621,0.00033608056,0.7736245,0.0004891765,0.00005025349,0.00031345323,0.02630532,0.0006282049,0.0024768836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99300086,0.0022299304,0.0008781007,0.0019161309,0.0016988214,0.000276157],"domain_scores_gemma":[0.9773111,0.0115756225,0.0028244376,0.002538152,0.0052999235,0.0004507943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039475174,0.0020233488,0.00097248005,0.008181075,0.00085891125,0.0022700203,0.0030005435,0.0020728745,0.003406782],"category_scores_gemma":[0.018711718,0.00053241145,0.0017006601,0.0034147806,0.00079465396,0.0027976593,0.001743358,0.0025242346,0.0034167448],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039564836,0.00065780664,0.013836767,0.001348865,0.00015588742,0.00064288545,0.00078295654,0.051770534,0.022070797,0.003925346,0.03497245,0.86944014],"study_design_scores_gemma":[0.00008026563,0.00016564659,0.005569999,0.00026458263,0.000076307995,0.000500647,0.0008384772,0.92949206,0.027664803,0.0086333165,0.026644,0.00006994153],"about_ca_topic_score_codex":0.011739465,"about_ca_topic_score_gemma":0.020573176,"teacher_disagreement_score":0.011739465,"about_ca_system_score_codex":0.0021786937,"about_ca_system_score_gemma":0.0026722853,"threshold_uncertainty_score":0.023342252},"labels":[],"label_agreement":null},{"id":"W3084309722","doi":"10.1145/3382494.3422172","title":"Profiling Developers Through the Lens of Technical Debt","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Technical debt; Code refactoring; Code smell; Commit; Debt; Profiling (computer programming); Coding (social sciences); Context (archaeology); Computer science; Software; Business; Software engineering; Software development; World Wide Web; Software quality; Finance; Database","score_opus":0.06876920546655389,"score_gpt":0.313855764819138,"score_spread":0.2450865593525841,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3084309722","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96811366,0.0010175767,0.016552877,0.0012015939,0.000028733877,0.00008052008,0.0065543517,0.00037285872,0.0060778568],"genre_scores_gemma":[0.9705202,0.0004638894,0.017231423,0.00022148067,0.000052198484,0.00012805502,0.009090287,0.00010355181,0.0021889557],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9960602,0.001314618,0.00042890658,0.0007702771,0.0011796973,0.000246342],"domain_scores_gemma":[0.9523943,0.018002119,0.01796679,0.0042497455,0.0056198225,0.0017672194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037670257,0.00042069572,0.00033979173,0.0094462875,0.0007086408,0.0018199907,0.0005491391,0.0006238635,0.0006873852],"category_scores_gemma":[0.036698263,0.0002963394,0.00020378125,0.0072411094,0.00045176988,0.0021565864,0.002212323,0.0008382729,0.0003710243],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006947776,0.00007046777,0.9188769,0.00022902695,0.000056687662,0.0003197546,0.0067491503,0.0008996053,0.003189481,0.0012606649,0.0052512507,0.063027576],"study_design_scores_gemma":[0.000010990582,0.00010543491,0.9506802,0.00017042452,0.000045500034,0.00074457686,0.004937395,0.009670344,0.0029858528,0.0037380971,0.026861554,0.000049550792],"about_ca_topic_score_codex":0.0046582883,"about_ca_topic_score_gemma":0.008249055,"teacher_disagreement_score":0.0094462875,"about_ca_system_score_codex":0.0006561039,"about_ca_system_score_gemma":0.0009190365,"threshold_uncertainty_score":0.019922197},"labels":[],"label_agreement":null},{"id":"W3085434149","doi":"10.1145/3387904.3389259","title":"The Secret Life of Commented-Out Source Code","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada; Concordia University of Edmonton","keywords":"Computer science; Code (set theory); Programming language; Source code; Code review; Program comprehension; Natural language; Natural (archaeology); KPI-driven code analysis; Software; Comprehension; Static program analysis; Software engineering; Software development; Natural language processing; Software system; History","score_opus":0.03565432937609865,"score_gpt":0.26941432173375945,"score_spread":0.23375999235766082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3085434149","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7177272,0.008826182,0.11458145,0.050541155,0.0021316973,0.00031510048,0.00085948734,0.0026215934,0.10239608],"genre_scores_gemma":[0.95415723,0.0021837489,0.01686497,0.0051417425,0.0009714741,0.00013351643,0.0004557063,0.0016724105,0.018419184],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.95409703,0.022944804,0.0017218433,0.0035584671,0.016080873,0.0015969401],"domain_scores_gemma":[0.65968865,0.19032821,0.05272445,0.044372655,0.04451703,0.008369042],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023211472,0.0007201895,0.00044077975,0.0037609998,0.0050211614,0.007904138,0.0014957642,0.0026092501,0.0059066857],"category_scores_gemma":[0.20676778,0.0009906556,0.0004854424,0.0023456048,0.011760035,0.014108234,0.006901933,0.004015424,0.0026355204],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006633807,0.00019850658,0.0713098,0.0017354058,0.00015854817,0.0028066908,0.40121472,0.00053808215,0.025378007,0.07887447,0.045618344,0.37150404],"study_design_scores_gemma":[0.000057340592,0.00054442894,0.059575308,0.005005814,0.00016472494,0.008673577,0.15229219,0.0034164842,0.01860928,0.08715767,0.66405475,0.00044839614],"about_ca_topic_score_codex":0.0014076615,"about_ca_topic_score_gemma":0.0019671153,"teacher_disagreement_score":0.023211472,"about_ca_system_score_codex":0.0029768746,"about_ca_system_score_gemma":0.0035919247,"threshold_uncertainty_score":0.12275541},"labels":[],"label_agreement":null},{"id":"W3085893373","doi":"10.1007/978-3-030-58923-3_14","title":"Architectural Technical Debt: A Grounded Theory","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Grounded theory; Computer science; Architectural pattern; Software engineering; Debt; Software; Software development; Management science; Qualitative research; Software design; Engineering; Business; Programming language; Finance; Sociology","score_opus":0.016722561261827036,"score_gpt":0.25364832393597747,"score_spread":0.23692576267415044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3085893373","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04681611,0.0034448358,0.48665577,0.023975762,0.0003471445,0.00062681537,0.0010977921,0.00028400225,0.43675175],"genre_scores_gemma":[0.8501956,0.0024935093,0.13087617,0.0011399859,0.00009239167,0.00054948445,0.0010479691,0.00021971024,0.013385124],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9951519,0.0030480041,0.00025776183,0.0005167082,0.0007617475,0.00026387523],"domain_scores_gemma":[0.9866377,0.009631632,0.0005856507,0.0016546233,0.0011992821,0.00029113496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006577335,0.00082156155,0.0006670531,0.0059839194,0.0037424946,0.0077776466,0.0027535132,0.0021350535,0.009806357],"category_scores_gemma":[0.013805249,0.00097545446,0.0007114181,0.007615053,0.017840624,0.016062649,0.0036710803,0.0042878943,0.0010218472],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000062042404,0.000015335827,0.00019199488,0.00005587178,0.0000034496154,0.000021194,0.002327214,0.00025395828,0.00003286466,0.9878755,0.0008410147,0.008375342],"study_design_scores_gemma":[0.000008074256,0.0000055788096,0.00020044949,0.00014672165,0.0000072876605,0.000030154994,0.0042540734,0.0012304868,0.00011250998,0.981296,0.012704035,0.0000045856573],"about_ca_topic_score_codex":0.007194731,"about_ca_topic_score_gemma":0.009318139,"teacher_disagreement_score":0.0104335835,"about_ca_system_score_codex":0.0104335835,"about_ca_system_score_gemma":0.006456937,"threshold_uncertainty_score":0.0757013},"labels":[],"label_agreement":null},{"id":"W3086435442","doi":"10.1145/3387904.3389262","title":"Investigating Near-Miss Micro-Clones in Evolving Software","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"clone (Java method); Commit; Computer science; Software; Code (set theory); Programming language; Source code; Base (topology); Software maintenance; Software system; Biology; Genetics; Mathematics; Database; Set (abstract data type)","score_opus":0.029206351388183737,"score_gpt":0.25868885856998086,"score_spread":0.22948250718179714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3086435442","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9889313,0.0003934319,0.010108693,0.000042669253,0.0000061084024,0.000028153694,0.000052019655,0.00008702453,0.0003506237],"genre_scores_gemma":[0.99173135,0.00015421437,0.0075851916,0.00002218751,0.000007784577,0.000020474801,0.00014655382,0.000022888327,0.00030938516],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99648917,0.00073398295,0.00028551914,0.00094336283,0.0013717165,0.00017627102],"domain_scores_gemma":[0.96061313,0.017903496,0.013294161,0.0025788012,0.004867235,0.0007431928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019204628,0.00026689004,0.0003835034,0.0028223912,0.0005212061,0.001071179,0.00057941483,0.0006401486,0.00036104958],"category_scores_gemma":[0.029729696,0.0002503323,0.00026985683,0.001934194,0.0005681201,0.0021630586,0.0009296622,0.0005160811,0.00012398962],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018540437,0.00010345123,0.8539032,0.00022497863,0.00012491127,0.0009708336,0.005773556,0.0044044834,0.015122505,0.0012148318,0.00033409023,0.11763777],"study_design_scores_gemma":[0.000011972593,0.00041869175,0.9200501,0.00007443604,0.00016593513,0.0030642648,0.0041112816,0.055780172,0.011076499,0.0024636926,0.0027277954,0.00005516376],"about_ca_topic_score_codex":0.0023059894,"about_ca_topic_score_gemma":0.0042219963,"teacher_disagreement_score":0.0028223912,"about_ca_system_score_codex":0.00043005252,"about_ca_system_score_gemma":0.00047354118,"threshold_uncertainty_score":0.010156512},"labels":[],"label_agreement":null},{"id":"W3086453443","doi":"10.1145/3372020.3391554","title":"UML Consistency Rules","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Junta de Comunidades de Castilla-La Mancha; Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada; European Commission; Ministerio de Ciencia, Innovación y Universidades","keywords":"UML tool; Applications of UML; Computer science; Unified Modeling Language; Class diagram; Consistency (knowledge bases); Programming language; Sequence diagram; Benchmark (surveying); Communication diagram; Software; Software engineering; Data mining; Artificial intelligence","score_opus":0.03400613067692527,"score_gpt":0.2596332412639361,"score_spread":0.22562711058701085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3086453443","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008045656,0.0031115855,0.80790377,0.0024960674,0.0015009984,0.004564205,0.02587759,0.024988623,0.121511504],"genre_scores_gemma":[0.0704442,0.0028991923,0.8328342,0.0022156008,0.0006127659,0.00454026,0.050853714,0.006771542,0.028828533],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9502126,0.013516336,0.010122061,0.005846755,0.018736122,0.0015660893],"domain_scores_gemma":[0.9312282,0.02450077,0.0056965454,0.017232325,0.020579956,0.00076213403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019069131,0.0028466405,0.0017899411,0.0118941795,0.0026545993,0.010034984,0.004569446,0.0037530062,0.022338124],"category_scores_gemma":[0.085061885,0.0017322924,0.003008844,0.006834346,0.0023373133,0.007754488,0.005141379,0.003556456,0.018672487],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030547485,0.0003471818,0.005463471,0.0025882684,0.0002709601,0.0009056628,0.0033367302,0.013581902,0.0065257614,0.3266546,0.14987148,0.49014848],"study_design_scores_gemma":[0.00007519924,0.00006436292,0.0008756776,0.0009993912,0.00011216052,0.000637296,0.00045755706,0.00998487,0.0054675494,0.0891302,0.89208037,0.00011529632],"about_ca_topic_score_codex":0.008505711,"about_ca_topic_score_gemma":0.006240437,"teacher_disagreement_score":0.022338124,"about_ca_system_score_codex":0.0032098424,"about_ca_system_score_gemma":0.007668726,"threshold_uncertainty_score":0.10084844},"labels":[],"label_agreement":null},{"id":"W3086619722","doi":"10.1145/1269900.1268837","title":"Introducing students to professional software construction","year":2007,"lang":"en","type":"article","venue":"ACM SIGCSE Bulletin","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Université du Québec à Montréal","keywords":"Documentation; Computer science; Software engineering; Key (lock); Software; Software construction; Software development; Software documentation; Software maintenance; Programming language","score_opus":0.011510398496954758,"score_gpt":0.29675010937146684,"score_spread":0.2852397108745121,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3086619722","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4481359,0.005088911,0.25374737,0.05457422,0.0031656404,0.0018952641,0.0007315603,0.005912689,0.22674835],"genre_scores_gemma":[0.62446254,0.005297601,0.16454966,0.0188823,0.0010126929,0.0011241765,0.00087005936,0.0006035899,0.18319733],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978283,0.0003893232,0.000118608754,0.0003224231,0.00061905256,0.0007222805],"domain_scores_gemma":[0.99297404,0.0009339035,0.000553892,0.00049098494,0.0015239399,0.0035232748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024130428,0.0010506555,0.0006810449,0.00092179765,0.0017429321,0.0047816373,0.0012747268,0.002135975,0.020084945],"category_scores_gemma":[0.007645157,0.00059953076,0.0007357509,0.0006251739,0.0021512737,0.0028956297,0.0058064535,0.0038518894,0.008115127],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032025858,0.008186553,0.032072034,0.0010369923,0.00002969122,0.0017202605,0.051743563,0.0027163369,0.03204444,0.14577173,0.18702814,0.5373301],"study_design_scores_gemma":[0.00008204331,0.0018290703,0.00977662,0.000724772,0.000030125913,0.003005584,0.009875158,0.004114488,0.014020024,0.06312872,0.89332455,0.000088822984],"about_ca_topic_score_codex":0.0010799237,"about_ca_topic_score_gemma":0.002004206,"teacher_disagreement_score":0.020084945,"about_ca_system_score_codex":0.0023650257,"about_ca_system_score_gemma":0.004525509,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3087274273","doi":"10.1109/icsme46990.2020.00062","title":"AOBTM: Adaptive Online Biterm Topic Modeling for Version Sensitive Short-texts Analysis","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Topic model; Inference; Online algorithm; Information retrieval; Mobile apps; Data mining; Data science; Machine learning; Artificial intelligence; World Wide Web; Algorithm","score_opus":0.07076359053981977,"score_gpt":0.31843126037649483,"score_spread":0.24766766983667504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3087274273","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01610684,0.0014206782,0.9757027,0.0003679521,0.00018653505,0.00027842136,0.0015898434,0.0037128513,0.00063408096],"genre_scores_gemma":[0.3170231,0.0023656131,0.6513303,0.0007614295,0.0011997912,0.0024916695,0.01574812,0.001088726,0.0079912385],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99661934,0.0013335824,0.000324853,0.0009861098,0.00053503504,0.00020110115],"domain_scores_gemma":[0.99326587,0.0044812197,0.0005388826,0.00070898526,0.00079183513,0.00021327507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036108831,0.0020343328,0.0018095899,0.003363117,0.00079649116,0.001707365,0.0028217381,0.002141345,0.0031492948],"category_scores_gemma":[0.01447369,0.0010116368,0.0025722599,0.0034884748,0.0007506984,0.0036376126,0.002677715,0.003219462,0.0033899494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001238495,0.0006430542,0.00782744,0.00095243496,0.00060233206,0.0005222582,0.0016562418,0.11895909,0.015319305,0.013975488,0.018451083,0.8198528],"study_design_scores_gemma":[0.0000507764,0.00008929477,0.0014330599,0.00003240813,0.000060088594,0.00010969376,0.00009901843,0.9792525,0.001593321,0.012505321,0.0047372794,0.00003733587],"about_ca_topic_score_codex":0.0056786435,"about_ca_topic_score_gemma":0.0077611054,"teacher_disagreement_score":0.0056786435,"about_ca_system_score_codex":0.0007811558,"about_ca_system_score_gemma":0.0013256978,"threshold_uncertainty_score":0.019096375},"labels":[],"label_agreement":null},{"id":"W3087926510","doi":"","title":"Toward the Automatic Classification of Self-Affirmed Refactoring","year":2020,"lang":"en","type":"article","venue":"RIT Scholar Works (Rochester Institute of Technology)","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Commit; Computer science; Code smell; Documentation; Java; Artificial intelligence; Machine learning; Quality (philosophy); Software engineering; Data mining; Software quality; Programming language; Database; Software; Software development","score_opus":0.041540164721141394,"score_gpt":0.26760586018599997,"score_spread":0.22606569546485858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3087926510","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.65113026,0.0025539994,0.32861733,0.0008693491,0.00030343106,0.00027670918,0.0032134755,0.0108759515,0.0021595745],"genre_scores_gemma":[0.8425328,0.0003568694,0.14642113,0.00017185044,0.00011688063,0.00014782452,0.007193339,0.00026334886,0.0027958727],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99788624,0.00040454394,0.00019204311,0.00072726124,0.0005792901,0.00021062282],"domain_scores_gemma":[0.9897413,0.004788203,0.0012428866,0.00084760005,0.003029642,0.00035036937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002622053,0.001172956,0.0010908532,0.005676964,0.0004568051,0.0013993437,0.0014560873,0.0013304688,0.0005465017],"category_scores_gemma":[0.009591589,0.0003038319,0.000812602,0.0019669684,0.00041501457,0.0015236543,0.0009544493,0.001887421,0.0010408926],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052500435,0.0008662687,0.14230324,0.00035041483,0.00015297173,0.0005772878,0.0006352628,0.021395557,0.028025122,0.0011521541,0.015101837,0.7889149],"study_design_scores_gemma":[0.000038648195,0.00018398219,0.027967496,0.0000901669,0.00007100686,0.0003310903,0.00026398004,0.9502839,0.013819184,0.002675473,0.004234332,0.000040832547],"about_ca_topic_score_codex":0.0071674427,"about_ca_topic_score_gemma":0.008171356,"teacher_disagreement_score":0.0071674427,"about_ca_system_score_codex":0.0007042348,"about_ca_system_score_gemma":0.0014795759,"threshold_uncertainty_score":0.014251411},"labels":[],"label_agreement":null},{"id":"W3088322718","doi":"10.1145/3372787.3390439","title":"On the detection of community smells using genetic programming-based ensemble classifier chain","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Code smell; Computer science; Classifier (UML); Machine learning; Closeness; Artificial intelligence; Software; Centrality; Software development; Software quality; Data science; Data mining; Knowledge management","score_opus":0.07273993122529661,"score_gpt":0.27905550458823986,"score_spread":0.20631557336294326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3088322718","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26095596,0.0008616274,0.7333696,0.0007314204,0.00010438878,0.00018562947,0.00022484119,0.0012892884,0.0022772413],"genre_scores_gemma":[0.8485445,0.00026673055,0.14691019,0.00036136512,0.000084220264,0.00014971448,0.0008413116,0.00011002888,0.002731941],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982603,0.0004948256,0.000095833966,0.0004749701,0.00045040465,0.0002236181],"domain_scores_gemma":[0.9921678,0.004736539,0.0005288521,0.0003469011,0.0019312329,0.0002887875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003051262,0.0013511003,0.001742097,0.0029493792,0.0010867588,0.00127297,0.0020252129,0.0021029997,0.0010185296],"category_scores_gemma":[0.008169119,0.00048173443,0.0012539822,0.0017813874,0.00073730765,0.0015557618,0.0013656549,0.0024212603,0.00042906508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023505319,0.00039518953,0.019334503,0.00008737855,0.00017023877,0.00021671384,0.00020556424,0.6720716,0.0036401115,0.0018533291,0.0020393603,0.299751],"study_design_scores_gemma":[0.0000026189805,0.000023571274,0.00032625426,0.0000042359015,0.000007860011,0.000011142332,0.00001191294,0.9986444,0.00036203905,0.000530269,0.00007256096,0.0000031537052],"about_ca_topic_score_codex":0.01824221,"about_ca_topic_score_gemma":0.012677566,"teacher_disagreement_score":0.01824221,"about_ca_system_score_codex":0.0012818609,"about_ca_system_score_gemma":0.0015939554,"threshold_uncertainty_score":0.03627205},"labels":[],"label_agreement":null},{"id":"W3089004665","doi":"10.1145/3387940.3392190","title":"Increasing the Trust In Refactoring Through Visualization","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Software engineering; Visualization; Process (computing); Software maintenance; Software development; Software; Programming language; Artificial intelligence","score_opus":0.049377824216383356,"score_gpt":0.3110515094955442,"score_spread":0.26167368527916085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089004665","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0814915,0.0018474383,0.8569588,0.006928951,0.00042011883,0.0004218589,0.00050283776,0.038431663,0.012996815],"genre_scores_gemma":[0.3995253,0.0011181606,0.5928369,0.00041675114,0.00014544206,0.00029522998,0.0004766736,0.0022476304,0.0029378776],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9934489,0.003907289,0.00046163783,0.00057295215,0.0013638514,0.000245429],"domain_scores_gemma":[0.93191475,0.039495423,0.0043437667,0.01320456,0.010007671,0.0010339296],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010963729,0.0020880974,0.0007488685,0.0025678326,0.00093733513,0.0042759706,0.0021144,0.0024699818,0.0056117047],"category_scores_gemma":[0.05975939,0.00089242595,0.0009224477,0.0014003549,0.0013313885,0.009939318,0.003572973,0.0035354916,0.0010581893],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014344046,0.00103349,0.019386707,0.0023303845,0.00022934708,0.0015736263,0.02582927,0.025779031,0.08833421,0.042326678,0.032295022,0.75944775],"study_design_scores_gemma":[0.0007065262,0.002077425,0.03182755,0.0034979705,0.0005811707,0.0029954372,0.0057822876,0.4377815,0.1528806,0.123922825,0.23662542,0.0013212828],"about_ca_topic_score_codex":0.002655789,"about_ca_topic_score_gemma":0.0021884972,"teacher_disagreement_score":0.010963729,"about_ca_system_score_codex":0.0009129647,"about_ca_system_score_gemma":0.0013021945,"threshold_uncertainty_score":0.057982385},"labels":[],"label_agreement":null},{"id":"W3089148641","doi":"10.1109/re48521.2020.00027","title":"Automated Recommendation of Templates for Legal Requirements","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Template; Context (archaeology); Precision and recall; Requirements engineering; Requirements elicitation; Recall; Information retrieval; Software engineering; Programming language; Psychology","score_opus":0.06913129860041883,"score_gpt":0.33498120117044167,"score_spread":0.26584990257002283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089148641","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046879932,0.00094244676,0.9283113,0.0012610307,0.000082918894,0.0018325454,0.0028492447,0.010317377,0.007523284],"genre_scores_gemma":[0.083738744,0.00045177905,0.90769005,0.00016481611,0.000025083342,0.00053577125,0.005654813,0.00040604683,0.0013330127],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98754346,0.005209867,0.00168012,0.0015211051,0.0037260165,0.0003194271],"domain_scores_gemma":[0.9478458,0.028016567,0.0057833404,0.008688748,0.009107283,0.00055827],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008923339,0.001561654,0.0011521456,0.008190706,0.0011474442,0.0027898757,0.0024609647,0.0018790129,0.0038001684],"category_scores_gemma":[0.059126794,0.0008497511,0.0020432442,0.003479274,0.0008397404,0.0031998404,0.0014467301,0.0013735141,0.0030622883],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031366345,0.0004908719,0.02408952,0.0022310263,0.00023836509,0.0014669005,0.0036862348,0.026909294,0.032645814,0.017673353,0.026406696,0.8638482],"study_design_scores_gemma":[0.00017456153,0.00042500964,0.017299602,0.0021213002,0.0005438958,0.0032816324,0.0042967084,0.6574871,0.12539051,0.037622083,0.15104702,0.00031063217],"about_ca_topic_score_codex":0.005849126,"about_ca_topic_score_gemma":0.009795159,"teacher_disagreement_score":0.008923339,"about_ca_system_score_codex":0.0015363715,"about_ca_system_score_gemma":0.0036684931,"threshold_uncertainty_score":0.04719168},"labels":[],"label_agreement":null},{"id":"W3089336588","doi":"10.1145/3387940.3392236","title":"Are Automatic Bug Report Summarizers Missing the Target?","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Annotation; Conversation; Sentence; Natural language processing; Software; Word (group theory); Interface (matter); Software bug; Gold standard (test); Artificial intelligence; Information retrieval; Programming language; Linguistics","score_opus":0.03465459722830837,"score_gpt":0.27059324695595893,"score_spread":0.23593864972765055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089336588","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17322123,0.006566205,0.6770674,0.009436687,0.00306035,0.001172581,0.013360222,0.1009443,0.015171013],"genre_scores_gemma":[0.4280232,0.0017934608,0.52896106,0.0018507921,0.0012773869,0.0013739889,0.01972697,0.007113554,0.009879578],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98298705,0.008044291,0.0020602266,0.0024219577,0.0038711268,0.0006153407],"domain_scores_gemma":[0.88447195,0.050364707,0.013569553,0.016758652,0.033223625,0.0016115687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017732287,0.0016439963,0.0012746832,0.004377145,0.0009342494,0.003704839,0.0019382106,0.001538161,0.004316378],"category_scores_gemma":[0.115393534,0.0008756429,0.0007423211,0.0023257423,0.000546327,0.0053401296,0.0018961383,0.0014080502,0.006099727],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078283536,0.0002254968,0.022966249,0.003297907,0.0002459064,0.0005970522,0.0068052006,0.0018507715,0.04345224,0.0038564089,0.0952395,0.82068044],"study_design_scores_gemma":[0.00037530396,0.0022461906,0.12625594,0.0029776338,0.0012300933,0.0027418484,0.0114531685,0.11710419,0.15346402,0.024610797,0.55669653,0.000844256],"about_ca_topic_score_codex":0.0020918914,"about_ca_topic_score_gemma":0.003205319,"teacher_disagreement_score":0.017732287,"about_ca_system_score_codex":0.00069934083,"about_ca_system_score_gemma":0.0015757364,"threshold_uncertainty_score":0.09377843},"labels":[],"label_agreement":null},{"id":"W3089375544","doi":"10.1145/3377811.3380380","title":"Context-aware in-process crowdworker recommendation","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Task (project management); Process (computing); Computer science; Context (archaeology); Software bug; Recommender system; Software; Focus (optics); Order (exchange); World Wide Web; Process management; Software engineering; Data science; Engineering; Operating system; Business","score_opus":0.03577391071284716,"score_gpt":0.2955875594301986,"score_spread":0.25981364871735146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089375544","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4253043,0.0028180175,0.54616874,0.0013062847,0.0003092474,0.0015400546,0.0010624882,0.009549581,0.01194123],"genre_scores_gemma":[0.8565914,0.0004444727,0.13784736,0.00021917379,0.000103140424,0.00044383775,0.000804243,0.00018112293,0.0033652356],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969952,0.0009632554,0.00016357696,0.0009021968,0.0007086247,0.00026717992],"domain_scores_gemma":[0.9916803,0.004317866,0.000815651,0.0011496934,0.0012370157,0.00079935935],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002988794,0.0013277417,0.0015276414,0.0016066035,0.0009784255,0.0016897066,0.001774616,0.0012591323,0.0027467678],"category_scores_gemma":[0.013506088,0.00053238025,0.0005418107,0.000808837,0.0003689894,0.0021100915,0.0018751767,0.0011729923,0.0015764086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002069641,0.002703339,0.089207806,0.0014097292,0.0003102968,0.0008207573,0.004714465,0.03586591,0.0451627,0.0028009077,0.014953838,0.7999807],"study_design_scores_gemma":[0.00040449345,0.0018539559,0.06948458,0.0004861897,0.00064846623,0.00093619793,0.0047274744,0.82507586,0.037106153,0.016642991,0.04231032,0.0003233816],"about_ca_topic_score_codex":0.0045591295,"about_ca_topic_score_gemma":0.009289558,"teacher_disagreement_score":0.0045591295,"about_ca_system_score_codex":0.00048650484,"about_ca_system_score_gemma":0.0013624745,"threshold_uncertainty_score":0.015806496},"labels":[],"label_agreement":null},{"id":"W3089622721","doi":"10.1016/j.infsof.2020.106440","title":"Special Section on the 2019 Symposium on Search-Based Software Engineering","year":2020,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Context (archaeology); Computer science; Data science; Diversity (politics); Measure (data warehouse); Data mining; Software; Machine learning; Geography","score_opus":0.010942623226967553,"score_gpt":0.21306865690337123,"score_spread":0.2021260336764037,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089622721","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00038795944,0.018939313,0.0016467769,0.0478391,0.8934578,0.00009241763,0.00039826112,0.00021202189,0.037026346],"genre_scores_gemma":[0.0029189936,0.01837064,0.0009492514,0.02116136,0.75264513,0.00012378689,0.0007895001,0.0003864397,0.20265491],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973008,0.00039363944,0.0002002447,0.00053372723,0.0012073284,0.0003643322],"domain_scores_gemma":[0.9899094,0.002536184,0.00066143495,0.00070995797,0.0037609984,0.0024220073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032211533,0.0026729908,0.0028919713,0.0051353313,0.002838351,0.0067208274,0.0027512226,0.0076961317,0.08942916],"category_scores_gemma":[0.0071581844,0.0007200595,0.0019058,0.0030125596,0.0012580132,0.004296262,0.003238985,0.0071861153,0.044742502],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000042292042,0.000037679372,0.00006984839,0.00011763832,0.000010528034,0.000039405255,0.000007735903,0.000040498267,0.00017789913,0.00086565793,0.9869611,0.0116296355],"study_design_scores_gemma":[0.00001787913,0.000050604205,0.00043912712,0.00018174486,0.000021340193,0.0000798473,0.000021545997,0.00025430147,0.00014999567,0.001345274,0.9974235,0.000014909254],"about_ca_topic_score_codex":0.0024586,"about_ca_topic_score_gemma":0.009439918,"teacher_disagreement_score":0.08942916,"about_ca_system_score_codex":0.0033347858,"about_ca_system_score_gemma":0.0024310295,"threshold_uncertainty_score":0.29917037},"labels":[],"label_agreement":null},{"id":"W3090054056","doi":"10.17705/1jais.00630","title":"(Re)considering the Concept of Literature Review Reproducibility","year":2020,"lang":"en","type":"article","venue":"Journal of the Association for Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; HEC Montréal; University of Waterloo","funders":"","keywords":"Terminology; Ambiguity; Transparency (behavior); Field (mathematics); Trustworthiness; Engineering ethics; Systematic review; Process (computing); Epistemology; Management science; Computer science; Data science; Political science; MEDLINE; Linguistics; Law; Engineering","score_opus":0.027570786993916616,"score_gpt":0.27701921086801834,"score_spread":0.2494484238741017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3090054056","genre_codex":"commentary","genre_gemma":"methods","domain_codex":"methods","domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008152631,0.051464897,0.19699268,0.645674,0.061529372,0.0042606676,0.00046434524,0.0012162111,0.030245263],"genre_scores_gemma":[0.3138202,0.024832346,0.30980042,0.27478337,0.04552516,0.018895885,0.0004955299,0.0019371763,0.009909818],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.17734091,0.52584904,0.13621637,0.03315977,0.12289364,0.0045403177],"domain_scores_gemma":[0.047334697,0.63561624,0.06766921,0.122589745,0.12252775,0.0042623505],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.6673602,0.0021737155,0.0049873856,0.018392274,0.0084155,0.026269501,0.010385138,0.020547869,0.004180205],"category_scores_gemma":[0.90384024,0.0030328957,0.004381024,0.012287069,0.04974565,0.044808153,0.020525511,0.026251037,0.0034787185],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057503965,0.00017806202,0.0076387343,0.03128271,0.002496783,0.0019139226,0.10658502,0.001244111,0.0026525839,0.46955648,0.16281222,0.2130643],"study_design_scores_gemma":[0.00047318943,0.0003330002,0.003166312,0.04593083,0.001260003,0.002038205,0.01610406,0.003119977,0.002746686,0.48837316,0.43584758,0.0006070075],"about_ca_topic_score_codex":0.0030423312,"about_ca_topic_score_gemma":0.0033020189,"teacher_disagreement_score":0.3326398,"about_ca_system_score_codex":0.012288511,"about_ca_system_score_gemma":0.043329414,"threshold_uncertainty_score":0.41020417},"labels":[],"label_agreement":null},{"id":"W3090250978","doi":"10.1145/3379597.3387456","title":"On the Relationship between User Churn and Software Issues","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Jargon; Software; Product (mathematics); Service (business); Competition (biology); Empirical research; World Wide Web; Software engineering; Business; Marketing","score_opus":0.08173328717630302,"score_gpt":0.30190875047011534,"score_spread":0.22017546329381232,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3090250978","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9958905,0.0003268795,0.0013141403,0.0004677566,0.000008576628,0.00002557765,0.00012953207,0.000023524168,0.0018135109],"genre_scores_gemma":[0.99900347,0.00009285662,0.00039258375,0.000054080618,0.000013205991,0.000013784331,0.00014386675,0.000008395888,0.00027769656],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.993761,0.0031881817,0.00051258394,0.00063134823,0.0012984659,0.0006083484],"domain_scores_gemma":[0.5073602,0.43222237,0.03916231,0.0047396165,0.009980671,0.0065348004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00643209,0.00038209022,0.0004505569,0.0025337238,0.0010623422,0.0018525858,0.0005740383,0.0012616183,0.0037065097],"category_scores_gemma":[0.10061428,0.00038697277,0.0007473735,0.0024794433,0.0009655373,0.002714983,0.0010966779,0.0028281438,0.00064291584],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019014053,0.0003372051,0.9881386,0.000044404274,0.000070545895,0.0000842277,0.00073329936,0.0009785147,0.00015434143,0.00030105625,0.00026743166,0.008700305],"study_design_scores_gemma":[0.000009448583,0.00029592338,0.97851163,0.000034919394,0.00007826367,0.00021146695,0.0011596442,0.018312436,0.00024619608,0.0006743428,0.0004342133,0.00003158645],"about_ca_topic_score_codex":0.0080207,"about_ca_topic_score_gemma":0.008476214,"teacher_disagreement_score":0.0080207,"about_ca_system_score_codex":0.001208152,"about_ca_system_score_gemma":0.0008097069,"threshold_uncertainty_score":0.03401655},"labels":[],"label_agreement":null},{"id":"W3090450084","doi":"10.1145/3377816.3381738","title":"Automatically predicting bug severity early in the development process","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Software bug; Computer science; Process (computing); Classifier (UML); Automation; Software; Predictive modelling; Software regression; Software engineering; Software development; Machine learning; Artificial intelligence; Engineering; Software quality; Programming language","score_opus":0.02609569107032494,"score_gpt":0.271727173197986,"score_spread":0.24563148212766106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3090450084","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81740314,0.0027535886,0.13779789,0.0010167997,0.0003463395,0.00048238068,0.010934619,0.025847876,0.0034173892],"genre_scores_gemma":[0.822506,0.0008323922,0.14802206,0.00019124545,0.00014620791,0.00019648006,0.025108889,0.0004580981,0.0025386491],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968182,0.0004835673,0.0004011048,0.0008601993,0.0012224241,0.00021454896],"domain_scores_gemma":[0.97565585,0.011126218,0.0038396842,0.0016332762,0.0069085513,0.00083644746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003442212,0.0018269897,0.0011497566,0.0072760913,0.0005715142,0.0016866206,0.0012504862,0.0011619234,0.00065857254],"category_scores_gemma":[0.020207422,0.00061875116,0.0008228128,0.003039073,0.00025062176,0.0021939112,0.00082756125,0.0015295125,0.0011566272],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036305786,0.0007562317,0.3775628,0.00072581664,0.00024477,0.00060581876,0.0006876188,0.03181193,0.018954128,0.00060323573,0.037243325,0.5304414],"study_design_scores_gemma":[0.0001176149,0.00071708654,0.16564499,0.0001776608,0.00024918982,0.00091979944,0.00042332298,0.78320867,0.02765305,0.0024324178,0.018329958,0.00012624003],"about_ca_topic_score_codex":0.011813015,"about_ca_topic_score_gemma":0.020468196,"teacher_disagreement_score":0.011813015,"about_ca_system_score_codex":0.00072814134,"about_ca_system_score_gemma":0.0017758412,"threshold_uncertainty_score":0.023488522},"labels":[],"label_agreement":null},{"id":"W3090595451","doi":"10.1145/3377811.3380433","title":"When APIs are intentionally bypassed","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Workaround; Application programming interface; Computer science; Software; Implementation; Software engineering; Operating system","score_opus":0.03422177729085972,"score_gpt":0.24864809826281478,"score_spread":0.21442632097195508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3090595451","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8902352,0.0007610137,0.07835487,0.004060517,0.0003339932,0.00049241126,0.00022146085,0.009307194,0.01623334],"genre_scores_gemma":[0.9488444,0.0004697034,0.03643387,0.0023357666,0.00014196968,0.00038180008,0.00044143575,0.0021957012,0.008755334],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9749341,0.007453431,0.0024153986,0.0026525194,0.009985536,0.0025590095],"domain_scores_gemma":[0.8861943,0.04778708,0.024265371,0.024305867,0.013786778,0.0036606495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014950547,0.0016354542,0.000793111,0.0034180034,0.002862209,0.00547935,0.002166729,0.0037024869,0.0029426573],"category_scores_gemma":[0.1415479,0.0017627933,0.0013707399,0.0020361557,0.0036854376,0.009944237,0.007216784,0.004116864,0.0014635128],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020755797,0.0007222472,0.27600932,0.0015754176,0.00052884186,0.018368969,0.1540366,0.0036873214,0.07232271,0.05137843,0.032905925,0.3863886],"study_design_scores_gemma":[0.00021011836,0.0022147875,0.15679911,0.003046065,0.0013107673,0.025723306,0.13484854,0.054724135,0.09130035,0.06978309,0.45880094,0.0012387723],"about_ca_topic_score_codex":0.0031068765,"about_ca_topic_score_gemma":0.00228909,"teacher_disagreement_score":0.014950547,"about_ca_system_score_codex":0.0018498772,"about_ca_system_score_gemma":0.0030598103,"threshold_uncertainty_score":0.07906699},"labels":[],"label_agreement":null},{"id":"W3090625562","doi":"10.1145/3377811.3380351","title":"Towards the use of the readily available tests from the release pipeline as performance tests","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Pipeline (software); Software performance testing; DevOps; Process (computing); Reliability engineering; Software; Software engineering; Non-regression testing; Software development; Software construction; Operating system; Engineering","score_opus":0.07500554734984366,"score_gpt":0.2557453517390077,"score_spread":0.18073980438916404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3090625562","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4712493,0.0066794315,0.4612663,0.0056085116,0.0009926005,0.0012411294,0.007283661,0.030432094,0.015246925],"genre_scores_gemma":[0.72888386,0.0009993962,0.25151172,0.0012229716,0.00038251793,0.0008374848,0.01164509,0.0028066395,0.0017102604],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.95179313,0.019119041,0.0037016987,0.0066678254,0.016770467,0.0019478364],"domain_scores_gemma":[0.592298,0.24002583,0.04821546,0.054989994,0.057202324,0.007268462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03638681,0.0027543763,0.0018526386,0.007972974,0.0011211295,0.0069825104,0.0036770115,0.0025175642,0.0018849918],"category_scores_gemma":[0.2742048,0.0009798062,0.0015341545,0.0046915943,0.001938517,0.010955681,0.0031179474,0.0054625687,0.0025280025],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015938109,0.0017934006,0.34864596,0.0020554892,0.0005400554,0.0011355262,0.0032365366,0.032675225,0.019341871,0.011366788,0.031004287,0.54661095],"study_design_scores_gemma":[0.0005188654,0.0035550687,0.3159721,0.0026470374,0.00057839917,0.0023805273,0.0040481505,0.48581102,0.048778147,0.050772663,0.08406211,0.0008759408],"about_ca_topic_score_codex":0.0070664147,"about_ca_topic_score_gemma":0.008312653,"teacher_disagreement_score":0.03638681,"about_ca_system_score_codex":0.0017557975,"about_ca_system_score_gemma":0.0034474651,"threshold_uncertainty_score":0.19243413},"labels":[],"label_agreement":null},{"id":"W3090668201","doi":"10.1145/3379597.3387467","title":"On the Prevalence, Impact, and Evolution of SQL Code Smells in Data-Intensive Systems","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code smell; Computer science; Code refactoring; SQL; Software quality; Code review; Code (set theory); Software engineering; Software; Database; Programming language; Software development","score_opus":0.049936564181284795,"score_gpt":0.2952276787439877,"score_spread":0.24529111456270292,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3090668201","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9981589,0.00024351827,0.0008365737,0.0001116474,0.000003618618,0.000013457984,0.000120402525,0.000026165428,0.0004857259],"genre_scores_gemma":[0.99846613,0.00013930233,0.0009084778,0.000029943096,0.0000071094228,0.00001245915,0.00023478601,0.000012176722,0.00018966012],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99271303,0.0017405555,0.0010352615,0.0012864392,0.0026802495,0.000544359],"domain_scores_gemma":[0.80669564,0.09905407,0.06871926,0.0057629007,0.016296202,0.0034718544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065880455,0.00033166702,0.00025370775,0.004328548,0.0005177855,0.0014491439,0.00053907896,0.00074416085,0.0009571398],"category_scores_gemma":[0.06238409,0.00035178202,0.0005058472,0.002727734,0.0010461679,0.0028926793,0.0015578928,0.0009470698,0.00027749073],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055535897,0.00006189961,0.9862125,0.00006189312,0.00004000673,0.00012446078,0.0016635694,0.00030180658,0.0007668695,0.0000904485,0.000099341654,0.010521609],"study_design_scores_gemma":[0.0000018249092,0.00009480779,0.9962852,0.000034598557,0.000021584963,0.00019697436,0.0011578538,0.0014259798,0.0004100345,0.00010611468,0.00025293868,0.000011924783],"about_ca_topic_score_codex":0.0037683037,"about_ca_topic_score_gemma":0.0064913565,"teacher_disagreement_score":0.0065880455,"about_ca_system_score_codex":0.0007334656,"about_ca_system_score_gemma":0.00080843904,"threshold_uncertainty_score":0.0348413},"labels":[],"label_agreement":null},{"id":"W3091074642","doi":"10.1145/3379597.3387479","title":"The Scent of Deep Learning Code","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Deep learning; Computer science; Artificial intelligence; Code (set theory); Machine learning; Code smell; Interpretability; Source code; Code review; Software quality; Static program analysis; Software; Software development; Programming language","score_opus":0.022824230042800025,"score_gpt":0.26057220948364535,"score_spread":0.23774797944084533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091074642","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8701785,0.0019007322,0.10556238,0.0046842787,0.00017873505,0.00006729978,0.0014052863,0.0051126494,0.010910206],"genre_scores_gemma":[0.9711902,0.0004839455,0.021462966,0.00051279325,0.0000628332,0.000042458345,0.000727378,0.0018946835,0.003622682],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99740297,0.000714416,0.0001422712,0.00036833182,0.0011979996,0.00017388212],"domain_scores_gemma":[0.96062183,0.02409771,0.0054020956,0.0047170673,0.0041377554,0.0010234589],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025020717,0.0003629987,0.00028978597,0.0020173064,0.0006269989,0.0025004132,0.0005356853,0.0007814994,0.0023772551],"category_scores_gemma":[0.03581958,0.0005260086,0.0004454704,0.0012354049,0.002386209,0.0035690896,0.001804,0.0017407152,0.0004940265],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017606241,0.00021980627,0.3203329,0.0014726932,0.00027972096,0.004874907,0.018917223,0.05099091,0.09487009,0.046841297,0.028041812,0.43139806],"study_design_scores_gemma":[0.00013484962,0.0012577958,0.40399837,0.0010850445,0.00026515606,0.007402939,0.007892771,0.2646433,0.086560026,0.12224343,0.10402015,0.0004961077],"about_ca_topic_score_codex":0.002194404,"about_ca_topic_score_gemma":0.0034362092,"teacher_disagreement_score":0.0025020717,"about_ca_system_score_codex":0.00078692805,"about_ca_system_score_gemma":0.0006364038,"threshold_uncertainty_score":0.01323235},"labels":[],"label_agreement":null},{"id":"W3091115553","doi":"10.1145/3377811.3380335","title":"Mitigating turnover with code review recommendation","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Workload; Context (archaeology); Code review; Recommender system; Replication (statistics); Code (set theory); Productivity; Work (physics); Knowledge management; Software; Data science; World Wide Web; Software development; Software quality; Engineering; Operating system","score_opus":0.03755239820781745,"score_gpt":0.2842022888922012,"score_spread":0.24664989068438378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091115553","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6243976,0.006009641,0.29898453,0.003852536,0.0007031489,0.0017077713,0.0013886571,0.042132817,0.020823171],"genre_scores_gemma":[0.7630954,0.0012518879,0.21375275,0.000978903,0.0003033359,0.00054966856,0.0026863066,0.0010160388,0.016365804],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99442434,0.0017606933,0.00040638587,0.0009974564,0.0020689806,0.00034218657],"domain_scores_gemma":[0.9618889,0.013876472,0.003797178,0.0060608312,0.012698222,0.0016784542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061118314,0.0010403289,0.0009449657,0.0015911426,0.0009991023,0.0020517695,0.0016094667,0.0011580834,0.0024987021],"category_scores_gemma":[0.043530077,0.0005555957,0.0005220427,0.0009627118,0.0002957842,0.0024272245,0.0012608687,0.0012621274,0.0028762645],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092045084,0.001024591,0.06867929,0.0008786994,0.00024251721,0.0003763572,0.0012764503,0.01514623,0.030793723,0.0008642259,0.043639418,0.8361581],"study_design_scores_gemma":[0.0005615322,0.003716781,0.090499155,0.0006984291,0.0010067205,0.0019471903,0.0025103805,0.7081621,0.063874125,0.005262081,0.12138394,0.00037755474],"about_ca_topic_score_codex":0.008439341,"about_ca_topic_score_gemma":0.02059796,"teacher_disagreement_score":0.008439341,"about_ca_system_score_codex":0.0008523222,"about_ca_system_score_gemma":0.0030106783,"threshold_uncertainty_score":0.032322824},"labels":[],"label_agreement":null},{"id":"W3091390234","doi":"","title":"UML Consistency Rules: a Case Study with Open-Source UML Models.","year":2020,"lang":"en","type":"article","venue":"Open Repository and Bibliography (University of Luxembourg)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Applications of UML; UML tool; Computer science; Class diagram; Unified Modeling Language; Consistency (knowledge bases); Programming language; Communication diagram; Sequence diagram; Benchmark (surveying); Systems Modeling Language; Software; Data mining; Software engineering; Artificial intelligence","score_opus":0.03619511658423043,"score_gpt":0.24237811123629552,"score_spread":0.20618299465206508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091390234","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90210813,0.0013452732,0.08493952,0.0013491841,0.00010169947,0.00135124,0.0013126427,0.0004959955,0.0069962186],"genre_scores_gemma":[0.86194533,0.0007883976,0.13154447,0.0003558528,0.000047073398,0.0006395947,0.0027658553,0.00029945755,0.0016140124],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9657336,0.01780706,0.003121461,0.0025021152,0.0098845465,0.000951234],"domain_scores_gemma":[0.8344033,0.1252967,0.012671084,0.011757043,0.014304104,0.001567806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025373152,0.0009949195,0.00070715055,0.005600014,0.002837871,0.0037016016,0.0036472464,0.0039398284,0.00073693576],"category_scores_gemma":[0.1000681,0.0006767217,0.0013258007,0.0069844946,0.0019991773,0.0045693815,0.0026188877,0.0021312372,0.00025929938],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010763657,0.0070407726,0.24637523,0.005282476,0.0006995057,0.05715592,0.14888585,0.056267034,0.02157901,0.039375074,0.015157672,0.4011052],"study_design_scores_gemma":[0.000660403,0.0030401808,0.17897339,0.004835632,0.0014328635,0.03456618,0.12854223,0.26314917,0.08840132,0.046785414,0.24873848,0.0008747057],"about_ca_topic_score_codex":0.013500162,"about_ca_topic_score_gemma":0.022194583,"teacher_disagreement_score":0.025373152,"about_ca_system_score_codex":0.003171113,"about_ca_system_score_gemma":0.0031196778,"threshold_uncertainty_score":0.13418764},"labels":[],"label_agreement":null},{"id":"W3091431866","doi":"10.1016/j.infsof.2020.106435","title":"Archetypes of delay: An analysis of online developer conversations on delayed work items in IBM Jazz","year":2020,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Agile software development; Computer science; IBM; Software engineering; Scope (computer science); Process management; Project management; Knowledge management; World Wide Web; Systems engineering; Engineering","score_opus":0.01878912662431273,"score_gpt":0.2615916783248603,"score_spread":0.24280255170054757,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091431866","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99454534,0.000182671,0.0014887955,0.00021870491,0.000016369175,0.000029912051,0.00023564906,0.000071230774,0.00321129],"genre_scores_gemma":[0.99735224,0.00009456527,0.0010670255,0.00005082039,0.000018658888,0.000034448178,0.00030530704,0.00005032325,0.0010265859],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99661934,0.0015314887,0.00023430168,0.00032052744,0.00092132785,0.00037300278],"domain_scores_gemma":[0.91272706,0.06913529,0.007987327,0.002261997,0.004855576,0.003032749],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036683818,0.00028515802,0.00032554878,0.0037853422,0.0021262902,0.0029704918,0.00087430264,0.0010601913,0.0021213607],"category_scores_gemma":[0.041356333,0.0003750131,0.00021959108,0.002904758,0.00095204293,0.0033392534,0.0023884303,0.0014835155,0.0006156869],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022863608,0.00062083855,0.59041685,0.00049411796,0.00012895359,0.0022657055,0.27100968,0.0017284809,0.013819733,0.00822232,0.0044210223,0.104585804],"study_design_scores_gemma":[0.00005921635,0.0005836017,0.6875589,0.00032979128,0.0001566351,0.0010989951,0.2581355,0.010791424,0.0046648863,0.004171484,0.03230332,0.00014627884],"about_ca_topic_score_codex":0.011891372,"about_ca_topic_score_gemma":0.014063851,"teacher_disagreement_score":0.011891372,"about_ca_system_score_codex":0.001873291,"about_ca_system_score_gemma":0.0019091155,"threshold_uncertainty_score":0.023644328},"labels":[],"label_agreement":null},{"id":"W3091435513","doi":"10.1145/3365438.3410953","title":"Leveraging natural-language requirements for deriving better acceptance criteria from models","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds National de la Recherche Luxembourg","keywords":"Computer science; Requirements analysis; Natural language; Software requirements specification; Non-functional testing; Requirements elicitation; Acceptance testing; Software engineering; Information system; System requirements specification; Software requirements; System requirements; Software; Non-functional requirement; Information retrieval; Software system; Software development; Artificial intelligence; Programming language; Software design; Engineering; Software construction","score_opus":0.06341128614539714,"score_gpt":0.316385960359451,"score_spread":0.25297467421405384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091435513","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0056016142,0.000063766136,0.9919063,0.00033496905,0.000013361583,0.00024616584,0.00013170215,0.0005238188,0.0011782665],"genre_scores_gemma":[0.08987807,0.0001657061,0.9071845,0.00024825236,0.000033882778,0.00073511194,0.00082170265,0.00046026197,0.00047260136],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.95297176,0.023521878,0.0036499095,0.002614122,0.016350063,0.0008922061],"domain_scores_gemma":[0.89012164,0.07612843,0.007283701,0.014550769,0.011398163,0.00051728735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022538848,0.0025796683,0.0015753463,0.008138712,0.0011005643,0.0046481974,0.0033228635,0.00339638,0.0022655677],"category_scores_gemma":[0.12826513,0.0023336622,0.005061804,0.0036502974,0.002923942,0.009957024,0.0051618535,0.005157685,0.001039414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025733688,0.00090536865,0.006151629,0.0019528508,0.00045459627,0.0031396141,0.008262073,0.3002441,0.034396686,0.4293639,0.003887389,0.21098457],"study_design_scores_gemma":[0.000072910356,0.00020638686,0.0008422706,0.0007250064,0.00018194426,0.00066120736,0.00064088614,0.7223336,0.019810652,0.23421943,0.020122582,0.0001831287],"about_ca_topic_score_codex":0.003063947,"about_ca_topic_score_gemma":0.0059776744,"teacher_disagreement_score":0.022538848,"about_ca_system_score_codex":0.002440506,"about_ca_system_score_gemma":0.0044408995,"threshold_uncertainty_score":0.11919826},"labels":[],"label_agreement":null},{"id":"W3091970108","doi":"10.1007/s10664-020-09851-6","title":"Publish or perish, but do not forget your software artifacts","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Deutsche Forschungsgemeinschaft; Deutscher Akademischer Austauschdienst","keywords":"Artifact (error); Publish or perish; Computer science; Publication; Context (archaeology); Replication (statistics); Data science; Software; Empirical research; Open science; Software engineering; World Wide Web; Publishing; Artificial intelligence; Political science","score_opus":0.06981769528065791,"score_gpt":0.2983438529525711,"score_spread":0.2285261576719132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091970108","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25101632,0.055414543,0.09041624,0.19397044,0.054083925,0.0015920047,0.051693123,0.011358517,0.29045492],"genre_scores_gemma":[0.7279933,0.029746272,0.073614255,0.021426616,0.0232576,0.0010488536,0.03157077,0.0056645363,0.085677885],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.95153445,0.01671341,0.005835456,0.0033123859,0.021276671,0.0013275832],"domain_scores_gemma":[0.5170412,0.22455345,0.06974537,0.104406245,0.07098224,0.013271527],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.039264724,0.0005856666,0.00087825913,0.01326257,0.0031891295,0.014657185,0.0018958541,0.0020797877,0.038384575],"category_scores_gemma":[0.33363926,0.0005147593,0.0011528082,0.02374813,0.0029277757,0.013517389,0.0059717502,0.0025810427,0.025200913],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042851432,0.00023368324,0.07632587,0.0066062603,0.0004945968,0.0011080619,0.0072791986,0.00063848204,0.0026832996,0.044304773,0.41170305,0.4481942],"study_design_scores_gemma":[0.00008438859,0.0001607789,0.028728208,0.0028115918,0.00015753704,0.0008668721,0.0028750626,0.0005024544,0.0021267787,0.036100093,0.92548233,0.000103776474],"about_ca_topic_score_codex":0.0012807144,"about_ca_topic_score_gemma":0.0021235896,"teacher_disagreement_score":0.96073526,"about_ca_system_score_codex":0.0018568309,"about_ca_system_score_gemma":0.0052235415,"threshold_uncertainty_score":0.20765418},"labels":[],"label_agreement":null},{"id":"W3092326336","doi":"","title":"ACM SIGSOFT Empirical Standards.","year":2020,"lang":"en","type":"article","venue":"Spiral (Imperial College London)","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Empirical research; Computer science; Quality (philosophy); Software; Data science; Engineering management; Engineering; Mathematics; Statistics","score_opus":0.037716144923894156,"score_gpt":0.30721679695565296,"score_spread":0.2695006520317588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3092326336","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025421013,0.004439478,0.21377806,0.024242563,0.008432055,0.0027884506,0.14720972,0.09361333,0.50295424],"genre_scores_gemma":[0.021155287,0.0071672285,0.26434365,0.006468714,0.002055564,0.00681107,0.3544915,0.0340765,0.30343053],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9724299,0.0058891587,0.005749528,0.0017353754,0.013445278,0.00075073313],"domain_scores_gemma":[0.8462946,0.03993996,0.006416689,0.039654586,0.06338476,0.004309488],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026690288,0.0021108827,0.0015905609,0.013388314,0.0023209145,0.011172072,0.005104686,0.0041646007,0.26537332],"category_scores_gemma":[0.1225559,0.002517344,0.0013828086,0.016504584,0.001833992,0.010032103,0.0048769168,0.006042591,0.27388734],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000088444285,0.00018250023,0.00093516294,0.00059383304,0.00001894535,0.000045688717,0.00025889627,0.00059081713,0.00030074856,0.021716028,0.77708393,0.19818498],"study_design_scores_gemma":[0.00005973739,0.000035114885,0.0012205445,0.0006487454,0.000015228451,0.000106743544,0.00017268796,0.0008724023,0.0002640409,0.023563445,0.9730102,0.000031016098],"about_ca_topic_score_codex":0.007170066,"about_ca_topic_score_gemma":0.0060187983,"teacher_disagreement_score":0.9733097,"about_ca_system_score_codex":0.0026637844,"about_ca_system_score_gemma":0.014211641,"threshold_uncertainty_score":0.8877622},"labels":[],"label_agreement":null},{"id":"W3092355216","doi":"10.1109/re48521.2020.00044","title":"Towards Queryable and Traceable Domain Models","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Rotation formalisms in three dimensions; Domain engineering; Domain model; Domain (mathematical analysis); TRACE (psycholinguistics); Software engineering; Domain analysis; Domain-specific language; Metric (unit); Artificial intelligence; Software development; Formalism (music); Software; Programming language; Domain knowledge; Component-based software engineering; Software construction; Engineering","score_opus":0.03431804165586858,"score_gpt":0.24622068333207078,"score_spread":0.2119026416762022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3092355216","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021325078,0.00020146016,0.9686763,0.00048288848,0.00004344778,0.0004656084,0.0014056013,0.005337771,0.0020618741],"genre_scores_gemma":[0.11296076,0.00030001372,0.8774647,0.00024281097,0.000021386722,0.0005880999,0.006287136,0.0008965097,0.0012386264],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9766381,0.010091597,0.0021133649,0.002315477,0.008131225,0.0007102892],"domain_scores_gemma":[0.9443328,0.027890006,0.0028510978,0.015882278,0.0082492875,0.00079457666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013192848,0.0015175209,0.0012595748,0.0041690622,0.0007146247,0.0072874385,0.0042506536,0.0032798592,0.0038183413],"category_scores_gemma":[0.074669555,0.001179684,0.0021796103,0.003434626,0.0014687511,0.011986619,0.0072115217,0.0043628295,0.0016590307],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014280784,0.0017307153,0.008253865,0.0017577336,0.0004116949,0.00087728206,0.003771598,0.38172436,0.032010794,0.1817858,0.010681367,0.37556675],"study_design_scores_gemma":[0.00009196109,0.00015275317,0.0005587932,0.00015296468,0.000054869397,0.00017019166,0.0005052374,0.8914197,0.015342051,0.07584001,0.015662642,0.000048828668],"about_ca_topic_score_codex":0.0068042446,"about_ca_topic_score_gemma":0.008394314,"teacher_disagreement_score":0.013192848,"about_ca_system_score_codex":0.002093133,"about_ca_system_score_gemma":0.0037954887,"threshold_uncertainty_score":0.06977129},"labels":[],"label_agreement":null},{"id":"W3093046707","doi":"10.1007/s10515-020-00277-4","title":"Understanding machine learning software defect predictions","year":2020,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":98,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Exploit; Machine learning; Software bug; Software; Artificial intelligence; Boosting (machine learning); Source code; Set (abstract data type); Data mining; Principal (computer security); Decision tree; Predictive modelling; Programming language","score_opus":0.04516996975354468,"score_gpt":0.24328561683835054,"score_spread":0.19811564708480586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093046707","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3721045,0.0010690435,0.61191314,0.0036128326,0.000100932724,0.00007930888,0.0009391448,0.002364038,0.007817058],"genre_scores_gemma":[0.9458072,0.00033874495,0.050567456,0.00016128046,0.00006580776,0.00003769566,0.0010597466,0.00008265714,0.0018794333],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9995121,0.0001478691,0.000033176933,0.000108592256,0.00014707797,0.000051314917],"domain_scores_gemma":[0.99416625,0.0044277227,0.00032386582,0.0003689728,0.00065235916,0.000060823237],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010344706,0.00051136443,0.0003845935,0.0011025813,0.000354617,0.0015244111,0.000988652,0.0009877929,0.0018422708],"category_scores_gemma":[0.011779897,0.00031674327,0.00037860379,0.00060140446,0.00044490446,0.0038100588,0.00068248896,0.0013383611,0.00048163207],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021073001,0.0005971022,0.047555048,0.00035885238,0.00013482115,0.00066164706,0.00089983316,0.30254227,0.009574272,0.103814214,0.016069014,0.5175823],"study_design_scores_gemma":[0.0000044698145,0.000021029919,0.0016109783,0.000021003041,0.000013895183,0.000050962368,0.00009547794,0.91986424,0.002065122,0.07461838,0.00162935,0.0000050678045],"about_ca_topic_score_codex":0.0037991197,"about_ca_topic_score_gemma":0.0041162185,"teacher_disagreement_score":0.0037991197,"about_ca_system_score_codex":0.00072326843,"about_ca_system_score_gemma":0.00065481546,"threshold_uncertainty_score":0.0075539947},"labels":[],"label_agreement":null},{"id":"W3093486972","doi":"10.1109/aire51212.2020.00011","title":"SENET: A Semantic Web for Supporting Automation of Software Engineering Tasks","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Vocabulary; Agile software development; Ontology; Domain (mathematical analysis); Semantic network; Set (abstract data type); Semantic Web; Glossary; Software engineering; Natural language; World Wide Web; Information retrieval; Natural language processing; Programming language","score_opus":0.01980360135938869,"score_gpt":0.26517510036656383,"score_spread":0.24537149900717514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093486972","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009478638,0.0007056664,0.927226,0.00080952747,0.00019012616,0.0004632677,0.0044728597,0.041633923,0.015020014],"genre_scores_gemma":[0.10305857,0.001914975,0.86162925,0.0007050284,0.00008701676,0.00060940604,0.01921036,0.0031606932,0.009624767],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985342,0.0003986487,0.00019496873,0.00020041627,0.000608483,0.00006319424],"domain_scores_gemma":[0.9977016,0.0009097825,0.00018672127,0.00064335985,0.0004255826,0.00013300369],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002460678,0.00087712216,0.0006286249,0.0031779865,0.001051115,0.0021116573,0.0014602684,0.001373921,0.0049355607],"category_scores_gemma":[0.005463007,0.0004935422,0.0009827806,0.0028542422,0.0011520227,0.008250026,0.002793735,0.0022297886,0.0023737564],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009732424,0.00059831684,0.0037611942,0.0022600095,0.0002106693,0.00112966,0.0017535562,0.042909864,0.030069435,0.2411297,0.085527025,0.58967745],"study_design_scores_gemma":[0.00014294799,0.00019260225,0.0019324529,0.0005088185,0.0001180198,0.0006511637,0.00047457137,0.32490957,0.02205628,0.14965466,0.49920583,0.00015307764],"about_ca_topic_score_codex":0.004696319,"about_ca_topic_score_gemma":0.00956663,"teacher_disagreement_score":0.0049355607,"about_ca_system_score_codex":0.0008501407,"about_ca_system_score_gemma":0.002089487,"threshold_uncertainty_score":0.016511083},"labels":[],"label_agreement":null},{"id":"W3093772117","doi":"10.1145/3382494.3422159","title":"Beyond Accuracy","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Analytics; Data science; Software analytics; Software; Empirical research; Software engineering; Software development; Software development process","score_opus":0.028221116142163948,"score_gpt":0.27756655203758357,"score_spread":0.24934543589541963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093772117","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011967847,0.058952488,0.55776817,0.13857351,0.0066134245,0.00034189405,0.009004879,0.004116624,0.21266119],"genre_scores_gemma":[0.53266054,0.050248228,0.3058059,0.029564548,0.014660693,0.00070799177,0.012500078,0.0036617965,0.05019018],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9543017,0.012987019,0.002997281,0.008738684,0.019532291,0.0014429723],"domain_scores_gemma":[0.78162485,0.1388524,0.008832473,0.035906464,0.032119356,0.002664471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025389893,0.001994704,0.0018550238,0.004561359,0.0024444168,0.016203906,0.0046982383,0.0046199127,0.037388463],"category_scores_gemma":[0.20077409,0.0007650932,0.001232965,0.0062161325,0.008005315,0.025725305,0.008727231,0.0067252214,0.02147551],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020315096,0.00006778447,0.009062421,0.0022941025,0.00020179314,0.00018540958,0.0008726967,0.008744025,0.0007889216,0.52700824,0.09628543,0.35428604],"study_design_scores_gemma":[0.000024100618,0.000057116522,0.0020826848,0.0020153555,0.00007868718,0.0004651441,0.00040934302,0.014980008,0.0013640863,0.7261722,0.25227302,0.00007810551],"about_ca_topic_score_codex":0.00640114,"about_ca_topic_score_gemma":0.0029990873,"teacher_disagreement_score":0.037388463,"about_ca_system_score_codex":0.00484119,"about_ca_system_score_gemma":0.0057374695,"threshold_uncertainty_score":0.13427621},"labels":[],"label_agreement":null},{"id":"W3094013987","doi":"10.1109/aire51212.2020.00013","title":"Traceability Network Analysis: A Case Study of Links in Issue Tracking Systems","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Traceability; TRACE (psycholinguistics); Computer science; Artifact (error); Perspective (graphical); Closure (psychology); Software; Property (philosophy); Resource (disambiguation); Requirements traceability; Process (computing); Link (geometry); Software engineering; Data mining; Data science; Software development; Artificial intelligence; Programming language","score_opus":0.04344609290700889,"score_gpt":0.306303241064092,"score_spread":0.2628571481570831,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3094013987","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81958985,0.00035917945,0.1669181,0.0011706969,0.00003327974,0.00042151182,0.0006334313,0.0005837945,0.010290099],"genre_scores_gemma":[0.92596215,0.00020756325,0.07060859,0.000055102737,0.000014001683,0.00015461388,0.000541052,0.00011215344,0.0023448083],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9956682,0.002138067,0.00026017567,0.00059420156,0.0010977507,0.00024161633],"domain_scores_gemma":[0.96458834,0.025230002,0.003169915,0.0028309552,0.0032339406,0.00094696076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005607741,0.00038182095,0.00029045367,0.0045133177,0.0029362456,0.0023028804,0.0011059706,0.0013155249,0.0018379915],"category_scores_gemma":[0.027146026,0.0002611022,0.00052050216,0.0061625694,0.0015628106,0.0053631663,0.0022545476,0.0013180269,0.00023188507],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051581685,0.0010183399,0.26138133,0.00096494274,0.00027059726,0.009955354,0.05846693,0.12299739,0.015480185,0.14840654,0.0077699344,0.3727726],"study_design_scores_gemma":[0.00010414596,0.00063479965,0.11924559,0.0004328256,0.00033161213,0.004113504,0.03717379,0.57627755,0.023826076,0.13410927,0.103535496,0.00021540765],"about_ca_topic_score_codex":0.010348815,"about_ca_topic_score_gemma":0.012728309,"teacher_disagreement_score":0.010348815,"about_ca_system_score_codex":0.001843957,"about_ca_system_score_gemma":0.0015030741,"threshold_uncertainty_score":0.029656887},"labels":[],"label_agreement":null},{"id":"W3094044696","doi":"10.1109/aire51212.2020.00017","title":"Analysis of Compatibility in Open Source Android Mobile Apps","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Compatibility (geochemistry); Computer science; Android (operating system); Commit; Software; World Wide Web; Software engineering; Database; Operating system; Engineering","score_opus":0.035915052183188294,"score_gpt":0.30612414514916986,"score_spread":0.27020909296598156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3094044696","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9943922,0.0005153038,0.0037918692,0.00009169343,0.000013750623,0.00007365266,0.00030515078,0.00020897757,0.000607368],"genre_scores_gemma":[0.9887656,0.00024897867,0.009101887,0.000044389333,0.000026997695,0.00008389774,0.00096706307,0.00007120914,0.00069006777],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99461126,0.0014338389,0.00069461373,0.0006522079,0.0023828563,0.00022527827],"domain_scores_gemma":[0.9230575,0.050311506,0.012594134,0.001841376,0.011393014,0.00080244156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031723646,0.000487736,0.0004955305,0.005716598,0.00047254667,0.0010647979,0.00038743662,0.00043223822,0.0003411041],"category_scores_gemma":[0.03362903,0.00029763818,0.00058570824,0.0022460448,0.00036060897,0.0014240966,0.00062167685,0.00062214216,0.00022925125],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080453604,0.0003770873,0.74100316,0.0011427322,0.00025365187,0.0021236371,0.009965605,0.0027291374,0.028263208,0.00043856882,0.0018786071,0.21102002],"study_design_scores_gemma":[0.00001320563,0.00074957794,0.94057477,0.00015674515,0.00017049498,0.0017677754,0.003234786,0.03906151,0.0102143595,0.00039009124,0.0035914194,0.000075363096],"about_ca_topic_score_codex":0.0039441776,"about_ca_topic_score_gemma":0.0057681687,"teacher_disagreement_score":0.005716598,"about_ca_system_score_codex":0.00047875088,"about_ca_system_score_gemma":0.0006486004,"threshold_uncertainty_score":0.016777277},"labels":[],"label_agreement":null},{"id":"W3094649757","doi":"10.1145/3432690","title":"Are Multi-Language Design Smells Fault-Prone? An Empirical Study","year":2021,"lang":"en","type":"preprint","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code smell; Computer science; Software design; Software engineering; Systems design; Software design pattern; Programming language; Software development; Software quality; Software","score_opus":0.233044127550916,"score_gpt":0.413159084720781,"score_spread":0.180114957169865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3094649757","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9987821,0.00017912682,0.0004728885,0.000083236424,0.000003274358,0.00002646178,0.00009103045,0.00001883063,0.00034297214],"genre_scores_gemma":[0.99879324,0.000121809964,0.00060199096,0.00004239833,0.0000068928675,0.000029641516,0.00017984293,0.000017789278,0.00020634006],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9918222,0.0023027244,0.001144472,0.0013480316,0.002787541,0.0005950461],"domain_scores_gemma":[0.77948684,0.1199407,0.07249368,0.008060135,0.016122278,0.0038963316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009739994,0.00052444934,0.0005042021,0.004022757,0.0008464536,0.0017086879,0.0012306712,0.0011103997,0.0017949765],"category_scores_gemma":[0.08003575,0.0005171061,0.0006251721,0.0024804368,0.0017114124,0.003263676,0.0017772177,0.001421885,0.00050005194],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015849505,0.00035640414,0.97663134,0.00019124208,0.00006737569,0.00052348664,0.007788373,0.0002084559,0.0006928384,0.00012025399,0.00042333323,0.012838407],"study_design_scores_gemma":[0.000013546546,0.00039254627,0.98376656,0.00013364888,0.00006682906,0.0008479844,0.009726038,0.0022524598,0.0008196118,0.00018651379,0.0017605912,0.000033665827],"about_ca_topic_score_codex":0.00252023,"about_ca_topic_score_gemma":0.0033670238,"teacher_disagreement_score":0.009739994,"about_ca_system_score_codex":0.001066954,"about_ca_system_score_gemma":0.0009199586,"threshold_uncertainty_score":0.051510632},"labels":[],"label_agreement":null},{"id":"W3095094616","doi":"10.1109/vissoft51673.2020.00011","title":"REM: Visualizing the Ripple Effect on Dependencies Using Metrics of Health","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Dependency graph; Computer science; Transitive relation; Dependency (UML); Visualization; Software; Graph; Theoretical computer science; Software visualization; Data mining; Transitive closure; Software development; Software engineering; Component-based software engineering; Programming language; Mathematics","score_opus":0.0963854050248794,"score_gpt":0.35506979807944333,"score_spread":0.2586843930545639,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3095094616","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1339027,0.002005371,0.795413,0.0021355005,0.00037455495,0.0003411744,0.007883455,0.042630844,0.0153133795],"genre_scores_gemma":[0.58953154,0.0014482011,0.39164063,0.00038066256,0.00010632301,0.00036353574,0.0057185576,0.005289162,0.0055213077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99908745,0.00026534862,0.00006958184,0.00013211794,0.00036077228,0.00008482103],"domain_scores_gemma":[0.9936342,0.0031588601,0.0008956815,0.00063911005,0.0013999295,0.00027205082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017608084,0.0016889699,0.0004877644,0.0047786683,0.00067804154,0.0017510467,0.0008035152,0.00088297686,0.004948899],"category_scores_gemma":[0.009562668,0.00044860633,0.0008295595,0.0024265388,0.00058455404,0.0027092153,0.001986094,0.0013054189,0.0006958095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089349045,0.00031820883,0.053491537,0.0028270704,0.0004476201,0.0024286916,0.010964556,0.11492912,0.06227188,0.06868467,0.08546532,0.5972779],"study_design_scores_gemma":[0.00014219123,0.00049733854,0.061969295,0.0009383947,0.00040794598,0.0018373076,0.0038475848,0.58482873,0.066760294,0.1243103,0.1539383,0.00052233343],"about_ca_topic_score_codex":0.008396394,"about_ca_topic_score_gemma":0.008705037,"teacher_disagreement_score":0.008396394,"about_ca_system_score_codex":0.0006459354,"about_ca_system_score_gemma":0.0010450364,"threshold_uncertainty_score":0.016695023},"labels":[],"label_agreement":null},{"id":"W3095187443","doi":"10.1109/vissoft51673.2020.00008","title":"Exploring Developer Preferences for Visualizing External Information Within Source Code Editors","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Source code; Software development; World Wide Web; Code (set theory); Software; Software engineering; Presentation (obstetrics); Data science; Programming language","score_opus":0.1491372542990679,"score_gpt":0.30221375720354904,"score_spread":0.15307650290448113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3095187443","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98544794,0.00011562074,0.011047018,0.00020261267,0.000008493199,0.000056849898,0.00006664897,0.00020788168,0.0028469444],"genre_scores_gemma":[0.9842429,0.00014120356,0.014284638,0.000065784785,0.0000048517786,0.000064549815,0.00016756909,0.00009298711,0.00093548314],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9948257,0.0034444607,0.00036292025,0.00033251263,0.00081628724,0.00021818858],"domain_scores_gemma":[0.8981162,0.08618139,0.005584528,0.0027768817,0.006124917,0.0012160883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010588433,0.00048598688,0.0002463505,0.0020180899,0.00044573838,0.003002127,0.00043803285,0.00073428,0.001966568],"category_scores_gemma":[0.072657734,0.000413079,0.00028934068,0.0008914524,0.00040227023,0.0026267162,0.0013498077,0.0005011915,0.0002905473],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018118782,0.0007525655,0.5212813,0.0014433042,0.00019867162,0.0015295226,0.15065603,0.0024249135,0.057829827,0.0033433023,0.0037860486,0.2549427],"study_design_scores_gemma":[0.00045844755,0.00308357,0.61687255,0.0016008605,0.0007428936,0.004196627,0.21825358,0.04716259,0.036451567,0.0045961225,0.06607668,0.00050453865],"about_ca_topic_score_codex":0.0018266498,"about_ca_topic_score_gemma":0.004714761,"teacher_disagreement_score":0.010588433,"about_ca_system_score_codex":0.00044773714,"about_ca_system_score_gemma":0.00050342985,"threshold_uncertainty_score":0.05599767},"labels":[],"label_agreement":null},{"id":"W3095662607","doi":"10.1109/icsme46990.2020.00058","title":"On the Impact of Multi-language Development in Machine Learning Frameworks","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Process (computing); Machine learning; Software engineering; Software development; Language acquisition; Software development process; Natural language processing; Programming language; Software; Mathematics education","score_opus":0.027056516848394756,"score_gpt":0.30460136757901873,"score_spread":0.27754485073062396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3095662607","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9808466,0.0016906674,0.01138066,0.0008284962,0.00007100462,0.00017268141,0.000099127516,0.0015864274,0.0033243005],"genre_scores_gemma":[0.982337,0.00045428987,0.014676073,0.00029980575,0.000046082227,0.000111289984,0.00028124903,0.0007340826,0.0010600439],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9500405,0.017204149,0.0038681272,0.0055021783,0.019904032,0.0034809855],"domain_scores_gemma":[0.55557245,0.31898093,0.061268236,0.023318524,0.033833653,0.007026248],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035087895,0.0010533128,0.00051951525,0.0044629383,0.0018758819,0.0047684084,0.0023620103,0.0013743463,0.001333347],"category_scores_gemma":[0.19334456,0.001061352,0.0011851517,0.0025293964,0.0027471452,0.0076614944,0.0053968243,0.0033827864,0.00046474722],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018621585,0.0015988899,0.5837843,0.0018573822,0.00051953475,0.0031058523,0.017143913,0.030862816,0.031272717,0.008264611,0.0054005836,0.31432727],"study_design_scores_gemma":[0.00033708647,0.0064015803,0.6694034,0.0012389519,0.0011598034,0.007392989,0.019303577,0.18631463,0.04986666,0.011288922,0.046529762,0.00076269946],"about_ca_topic_score_codex":0.00505661,"about_ca_topic_score_gemma":0.0061481474,"teacher_disagreement_score":0.035087895,"about_ca_system_score_codex":0.0027661815,"about_ca_system_score_gemma":0.0037213832,"threshold_uncertainty_score":0.18556476},"labels":[],"label_agreement":null},{"id":"W3095737939","doi":"10.1109/modre51215.2020.00010","title":"Detecting Emergent Behavior in Scenario-Based Specifications using a Probabilistic Model","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Probabilistic logic; Software engineering; Requirements elicitation; Software; Data mining; Artificial intelligence; Requirements analysis; Programming language","score_opus":0.17913984695598204,"score_gpt":0.31537680350087777,"score_spread":0.13623695654489573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3095737939","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041000832,0.000042798853,0.9569159,0.00013563495,0.000007635598,0.00011685464,0.00017408119,0.00065725035,0.0009490105],"genre_scores_gemma":[0.601624,0.00018003445,0.39597735,0.00008352816,0.000020181165,0.00045534622,0.0006098353,0.00013131904,0.00091844186],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951278,0.0018733789,0.00041577689,0.0006067525,0.001718273,0.00025802932],"domain_scores_gemma":[0.9827498,0.012065413,0.0023893905,0.0012130785,0.0013637239,0.000218563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036862693,0.0009488473,0.0005647369,0.0017003759,0.00056242844,0.0016873435,0.001292333,0.0012562667,0.0012357915],"category_scores_gemma":[0.015938893,0.0007771353,0.0017380168,0.0009077486,0.0014919797,0.0033437966,0.0012629433,0.0012448942,0.0003693679],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017388204,0.00011043629,0.009044926,0.00019551483,0.00010828827,0.001136788,0.000952756,0.8534475,0.015106019,0.096288435,0.00046107962,0.022974364],"study_design_scores_gemma":[0.0000107180595,0.000058445035,0.0005693516,0.000014810684,0.00002902522,0.00015376012,0.000056007204,0.97139984,0.0026307253,0.024198852,0.0008519913,0.000026420052],"about_ca_topic_score_codex":0.004592394,"about_ca_topic_score_gemma":0.0052717356,"teacher_disagreement_score":0.004592394,"about_ca_system_score_codex":0.001025503,"about_ca_system_score_gemma":0.0014342613,"threshold_uncertainty_score":0.01949507},"labels":[],"label_agreement":null},{"id":"W3095950387","doi":"10.1145/3417990.3421385","title":"DoMoBOT","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Domain (mathematical analysis); Domain model; Domain engineering; Domain analysis; Tracing; Artificial intelligence; Software engineering; Domain-specific language; Extractor; Software; Machine learning; Programming language; Domain knowledge; Software development; Component-based software engineering; Software construction","score_opus":0.03179134385973498,"score_gpt":0.25997537570294404,"score_spread":0.22818403184320907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3095950387","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041755647,0.00084593456,0.8031051,0.00065464806,0.00029361193,0.0010662414,0.006756343,0.14461678,0.038485777],"genre_scores_gemma":[0.046514485,0.0012967135,0.8696883,0.0012267487,0.000064916785,0.0013376191,0.02850203,0.022012975,0.029356247],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972824,0.0006633357,0.00026980703,0.0005901741,0.0009720035,0.00022228481],"domain_scores_gemma":[0.9963463,0.0018651468,0.00023983138,0.00090847735,0.00047743588,0.00016286345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026137375,0.0019969484,0.0011766421,0.002959592,0.0010050194,0.0039537842,0.0039511602,0.0022684291,0.02567292],"category_scores_gemma":[0.01140616,0.0015317635,0.002439263,0.0014481057,0.0013458162,0.005617226,0.0064812433,0.003458682,0.015795223],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010570618,0.00070531765,0.0039063855,0.0056477985,0.00026691632,0.0013104569,0.0020406824,0.0253685,0.01951948,0.19495517,0.25194728,0.493275],"study_design_scores_gemma":[0.00016605137,0.00011865187,0.00083042786,0.00051737373,0.000049966377,0.0008429374,0.00018769808,0.10908345,0.012369251,0.053577844,0.8221598,0.0000965242],"about_ca_topic_score_codex":0.005204317,"about_ca_topic_score_gemma":0.0084848665,"teacher_disagreement_score":0.02567292,"about_ca_system_score_codex":0.0018523319,"about_ca_system_score_gemma":0.0029554863,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3096162789","doi":"10.1007/s10664-020-09874-z","title":"A feature location approach for mapping application features extracted from crowd-based screencasts to source code","year":2020,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Concordia University","funders":"","keywords":"Computer science; Source code; Codebase; Set (abstract data type); Program comprehension; Software; Workflow; Code (set theory); Information retrieval; Multimedia; World Wide Web; Human–computer interaction; Database; Software system; Programming language","score_opus":0.03464601524693783,"score_gpt":0.2754876655325143,"score_spread":0.24084165028557647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096162789","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11183899,0.00033698356,0.8720528,0.00019313321,0.00013047237,0.00022448413,0.0023087158,0.0077651152,0.005149411],"genre_scores_gemma":[0.6687675,0.00018033282,0.32152003,0.000062972445,0.00008685599,0.00028520063,0.0027853923,0.0002760505,0.0060356217],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9994062,0.000083624225,0.000024173358,0.0002121068,0.00019099619,0.00008285067],"domain_scores_gemma":[0.9987962,0.0003151611,0.0001440806,0.00022562545,0.00044435685,0.000074491705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000439507,0.0007582382,0.0006053616,0.0042827753,0.0005083641,0.00092437014,0.0007670635,0.0009168015,0.0024411064],"category_scores_gemma":[0.0030092201,0.00025070252,0.0006225725,0.0026375942,0.0003323415,0.0010931528,0.0015770827,0.00059852784,0.0020721692],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007864829,0.0005829409,0.02861005,0.00030823867,0.00020723534,0.0008516751,0.0016142591,0.02314644,0.097427525,0.0055356016,0.015761767,0.8251678],"study_design_scores_gemma":[0.00009001906,0.00047225328,0.06280711,0.00007683021,0.00021642569,0.000962052,0.0019682073,0.85438555,0.04650876,0.012010929,0.020331426,0.00017047478],"about_ca_topic_score_codex":0.010033718,"about_ca_topic_score_gemma":0.014169695,"teacher_disagreement_score":0.010033718,"about_ca_system_score_codex":0.00038660757,"about_ca_system_score_gemma":0.0007569639,"threshold_uncertainty_score":0.019950628},"labels":[],"label_agreement":null},{"id":"W3096440690","doi":"10.1109/icsme46990.2020.00052","title":"Characterizing Task-Relevant Information in Natural Language Software Artifacts","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Natural language; Relevance (law); Task (project management); Consistency (knowledge bases); Semantics (computer science); Software documentation; Artifact (error); Natural language processing; Software; Documentation; Software development; Key (lock); Frame (networking); Artificial intelligence; Information retrieval; Programming language; Software development process","score_opus":0.011646656451223971,"score_gpt":0.2386682515069486,"score_spread":0.2270215950557246,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096440690","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9423931,0.00031979222,0.05444155,0.00020566996,0.000032101612,0.0004140809,0.0005217351,0.00037399086,0.0012978897],"genre_scores_gemma":[0.9205474,0.00022746182,0.07643785,0.00018805866,0.000040944593,0.00047764016,0.0014998406,0.00011843981,0.00046246438],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9927248,0.004182986,0.00064133113,0.0010723675,0.0011608766,0.00021759246],"domain_scores_gemma":[0.8779576,0.101002164,0.011860878,0.0024958653,0.0060193418,0.0006641203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005385827,0.0007433369,0.00041709366,0.002309017,0.000730281,0.0019759312,0.00069775677,0.0012409928,0.0009245511],"category_scores_gemma":[0.06523847,0.00036682974,0.000441965,0.0013111862,0.00089343754,0.0025084359,0.00092236244,0.0007403222,0.00029365573],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036599569,0.0024402365,0.13272019,0.008658414,0.00028218128,0.003955949,0.1195852,0.008654765,0.4790673,0.008409657,0.0055285734,0.22703753],"study_design_scores_gemma":[0.00071052916,0.0066304724,0.51479405,0.001489416,0.0008118201,0.004939308,0.05200456,0.11918672,0.21525767,0.02810097,0.05538279,0.0006916887],"about_ca_topic_score_codex":0.0015601476,"about_ca_topic_score_gemma":0.0022296973,"teacher_disagreement_score":0.005385827,"about_ca_system_score_codex":0.0009187589,"about_ca_system_score_gemma":0.00087712094,"threshold_uncertainty_score":0.028483331},"labels":[],"label_agreement":null},{"id":"W3096566855","doi":"10.1109/icsme46990.2020.00030","title":"A Fine-Grained Analysis on the Inconsistent Changes in Code Clones","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"clone (Java method); Java; Computer science; Software maintenance; Software bug; Software evolution; Source code; Code (set theory); Programming language; Software system; Software development; Software; Biology; Software construction; Genetics; Set (abstract data type); Gene","score_opus":0.046917457824916645,"score_gpt":0.27024986997261924,"score_spread":0.2233324121477026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096566855","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9754302,0.00034458903,0.022710996,0.00005874195,0.000006837421,0.00011220494,0.00030872977,0.00020505495,0.00082262955],"genre_scores_gemma":[0.9826888,0.00009968414,0.016330628,0.000020932426,0.0000069713196,0.000039956045,0.000421438,0.000032918895,0.00035860855],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9969528,0.000448347,0.00033270978,0.00069342804,0.0013701726,0.00020245969],"domain_scores_gemma":[0.9499962,0.023785545,0.013725308,0.0051575378,0.0065245028,0.00081089634],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021488992,0.00027333648,0.00035460637,0.004922207,0.00058840885,0.00094366865,0.00047858665,0.00045260743,0.000772899],"category_scores_gemma":[0.027496098,0.00025880095,0.00053568993,0.003091621,0.00057510287,0.001515774,0.0008297248,0.0005833508,0.00012354362],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002449319,0.00012351578,0.86957705,0.00021736094,0.00012479846,0.00086918514,0.0022905462,0.005826986,0.02343124,0.0021909624,0.00029420183,0.09480928],"study_design_scores_gemma":[0.000013497088,0.00030270673,0.94874233,0.00005548115,0.0001438852,0.0015475964,0.001032384,0.035231818,0.007906181,0.0026962815,0.0022864328,0.000041447787],"about_ca_topic_score_codex":0.003025514,"about_ca_topic_score_gemma":0.0041126236,"teacher_disagreement_score":0.004922207,"about_ca_system_score_codex":0.0005068405,"about_ca_system_score_gemma":0.00079746783,"threshold_uncertainty_score":0.011364579},"labels":[],"label_agreement":null},{"id":"W3096688843","doi":"10.1109/icsme46990.2020.00020","title":"Analysis of Modern Release Engineering Topics : – A Large-Scale Study using StackOverflow –","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Software deployment; Computer science; Software engineering; Merge (version control); Debugging; World Wide Web","score_opus":0.02667465741002027,"score_gpt":0.2706080302458039,"score_spread":0.24393337283578362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096688843","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99353737,0.0006494395,0.0029612232,0.00021826915,0.0000100907,0.00009017799,0.00062551687,0.000065225955,0.0018427406],"genre_scores_gemma":[0.99180436,0.0008200258,0.0034310536,0.000102622114,0.000045017452,0.00018308018,0.0020198082,0.000104513914,0.0014894886],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9954419,0.0014456435,0.000467487,0.0007919331,0.0014991161,0.00035388503],"domain_scores_gemma":[0.90881085,0.06339881,0.015139841,0.00302223,0.007551051,0.0020771767],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0069343024,0.00039648407,0.00033789294,0.0061274725,0.0012742335,0.002333963,0.00058302394,0.00060139527,0.0014129067],"category_scores_gemma":[0.03575559,0.0003515524,0.00055755395,0.0055553913,0.00087531586,0.005296634,0.0015797631,0.0011214323,0.0005111675],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002755939,0.0005309927,0.72958857,0.0012746546,0.00013421297,0.0006569879,0.09446016,0.0010273467,0.005934358,0.002641923,0.006097329,0.15737796],"study_design_scores_gemma":[0.000015125537,0.00025119883,0.9102628,0.0003527094,0.00008173887,0.0005236961,0.05337451,0.0051064542,0.0027992683,0.0009571014,0.026190814,0.0000843996],"about_ca_topic_score_codex":0.0042837677,"about_ca_topic_score_gemma":0.00540128,"teacher_disagreement_score":0.9938725,"about_ca_system_score_codex":0.0014658727,"about_ca_system_score_gemma":0.0012334754,"threshold_uncertainty_score":0.036672533},"labels":[],"label_agreement":null},{"id":"W3096761425","doi":"10.1109/icsme46990.2020.00075","title":"SiblingClassTestDetector: Finding Untested Sibling Functions","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Test suite; Implementation; Suite; Set (abstract data type); Open source; Software engineering; Test (biology); Test case; Programming language; Software; Machine learning","score_opus":0.06061521721568394,"score_gpt":0.27797323455558526,"score_spread":0.21735801733990132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096761425","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2694947,0.0007885393,0.47082585,0.00080313045,0.0002523305,0.00096705294,0.01139062,0.23810361,0.0073741605],"genre_scores_gemma":[0.569178,0.0002704474,0.36709115,0.0007023902,0.00007842581,0.0010729432,0.02588287,0.02878779,0.00693606],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9919125,0.0018483169,0.0008120086,0.0017205959,0.003010056,0.00069654937],"domain_scores_gemma":[0.95630336,0.027310653,0.004650866,0.005531624,0.0050355783,0.0011678339],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069218553,0.0027796333,0.0010720198,0.004281655,0.00086176896,0.0020252282,0.0035526475,0.0020843917,0.008570352],"category_scores_gemma":[0.036617808,0.0011351823,0.001392502,0.0019006352,0.0013477292,0.0034749596,0.0028569715,0.0011741136,0.0030489361],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024525116,0.0011889263,0.17775075,0.0033305967,0.00048189232,0.0036573939,0.003695313,0.016153298,0.07718474,0.012113097,0.18257134,0.51942015],"study_design_scores_gemma":[0.0011645472,0.0020362798,0.07923373,0.0010663996,0.0004964585,0.004812784,0.002353379,0.34132585,0.34386533,0.01625656,0.20667723,0.00071142515],"about_ca_topic_score_codex":0.0052781785,"about_ca_topic_score_gemma":0.007924783,"teacher_disagreement_score":0.008570352,"about_ca_system_score_codex":0.00093708513,"about_ca_system_score_gemma":0.003138623,"threshold_uncertainty_score":0.03660673},"labels":[],"label_agreement":null},{"id":"W3096935998","doi":"10.48550/arxiv.2004.08378","title":"On Using Stack Overflow Comment-Edit Pairs to recommend code maintenance changes","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Commit; Code (set theory); Source code; Stack (abstract data type); Programming language; Code smell; Code review; Point (geometry); Data mining; Static program analysis; Software engineering; Software; Database; Software development; Set (abstract data type); Software quality","score_opus":0.15031486902825433,"score_gpt":0.2347062665468398,"score_spread":0.08439139751858546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096935998","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.828057,0.004175207,0.1235399,0.001164545,0.00043717795,0.0009271149,0.020909611,0.0160237,0.0047657653],"genre_scores_gemma":[0.6886393,0.00082852645,0.25310096,0.00047046618,0.0002575293,0.00044992016,0.051126733,0.00085664575,0.0042699454],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99567765,0.0008806966,0.0005363494,0.0011653806,0.0014869985,0.00025285096],"domain_scores_gemma":[0.96535134,0.019477796,0.0050230185,0.002589221,0.0066205533,0.0009380059],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037042673,0.0017993135,0.00092749327,0.0153444195,0.0010426143,0.002020323,0.0013474181,0.0018653052,0.0012956471],"category_scores_gemma":[0.0324061,0.0005628298,0.00071423425,0.0050256806,0.00056715024,0.0043822117,0.0018490508,0.0015424518,0.0018345419],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017510125,0.0009193406,0.4155551,0.0020056942,0.00052285125,0.0016502017,0.003655861,0.0145502575,0.019780919,0.0019223978,0.042507034,0.49517938],"study_design_scores_gemma":[0.00033079874,0.0016815503,0.279988,0.0008778483,0.00063697307,0.004268048,0.00523474,0.5964367,0.0394612,0.0058781835,0.06474968,0.00045634163],"about_ca_topic_score_codex":0.006353598,"about_ca_topic_score_gemma":0.01551036,"teacher_disagreement_score":0.0153444195,"about_ca_system_score_codex":0.0005998305,"about_ca_system_score_gemma":0.0015066563,"threshold_uncertainty_score":0.019590259},"labels":[],"label_agreement":null},{"id":"W3096975346","doi":"10.1007/s11390-020-9668-1","title":"Machine Learning Techniques for Software Maintainability Prediction: Accuracy Analysis","year":2020,"lang":"en","type":"article","venue":"Journal of Computer Science and Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Maintainability; Machine learning; Computer science; Artificial intelligence; Support vector machine; Software; Artificial neural network; Data mining; Reliability engineering; Software engineering; Engineering","score_opus":0.01323511758263632,"score_gpt":0.27447270302207677,"score_spread":0.2612375854394404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096975346","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3528957,0.005774235,0.63393146,0.0011079311,0.00019992815,0.00008733555,0.0008615651,0.0020129099,0.003128879],"genre_scores_gemma":[0.91864663,0.0010913167,0.07744608,0.00010052419,0.00016053644,0.000050282302,0.00095879193,0.000082897546,0.0014630033],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968141,0.0010006753,0.00036055341,0.00037755005,0.0012730489,0.00017419826],"domain_scores_gemma":[0.96905524,0.022898028,0.0015403349,0.0021933427,0.004128047,0.00018506216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058350903,0.0009006947,0.0011268379,0.0032833512,0.00046466925,0.001277253,0.0012903983,0.0012640746,0.0009520157],"category_scores_gemma":[0.030476063,0.00027999093,0.0008625113,0.002144604,0.00032892247,0.0017179936,0.00072198926,0.0017900007,0.00057200575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039590264,0.00040443614,0.06254731,0.00022211281,0.00037386167,0.00007539256,0.00010248079,0.27452803,0.003667108,0.0017840663,0.0030064704,0.6528929],"study_design_scores_gemma":[0.000008985484,0.000096444084,0.0054976926,0.000025065325,0.000056961515,0.000047997524,0.00001573656,0.990274,0.0022151053,0.0014472436,0.00030491236,0.000009854648],"about_ca_topic_score_codex":0.005471905,"about_ca_topic_score_gemma":0.0047158552,"teacher_disagreement_score":0.0058350903,"about_ca_system_score_codex":0.0007098483,"about_ca_system_score_gemma":0.000717418,"threshold_uncertainty_score":0.030859232},"labels":[],"label_agreement":null},{"id":"W3097665816","doi":"10.1109/icsme46990.2020.00038","title":"Studying Software Developer Expertise and Contributions in Stack Overflow and GitHub","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; JavaScript; Key (lock); Context (archaeology); Software; Subject-matter expert; Knowledge management; Exploratory research; Field (mathematics); World Wide Web; Data science; Software development; Software engineering; Expert system; Artificial intelligence; Computer security","score_opus":0.031047070550848842,"score_gpt":0.2708157261517995,"score_spread":0.23976865560095065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3097665816","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9965659,0.00020705204,0.00094679784,0.00019953973,0.000006959548,0.000020037698,0.000017085102,0.000009195159,0.0020274227],"genre_scores_gemma":[0.9980715,0.00020753955,0.0009143096,0.00006813301,0.0000096819585,0.000025411102,0.00003322961,0.000010393854,0.00065989495],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9932221,0.003063641,0.00048617617,0.00073303404,0.0017095531,0.0007854428],"domain_scores_gemma":[0.931811,0.044652387,0.011417226,0.0014611324,0.0058947746,0.0047634286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00868684,0.00037453187,0.00039209804,0.0036890074,0.0012729125,0.0021747653,0.0006575453,0.0008512594,0.0015100633],"category_scores_gemma":[0.05022949,0.00035803026,0.00027316558,0.0016877869,0.0016747446,0.003708308,0.0031390884,0.0008288907,0.00028833663],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014118115,0.00032205312,0.46834224,0.00054900185,0.00006109696,0.0014498563,0.46581036,0.00049788225,0.0029983402,0.0009098331,0.0010262253,0.057891916],"study_design_scores_gemma":[0.000024150704,0.0004232691,0.6235643,0.0005637332,0.000065983404,0.001755734,0.35714042,0.002058698,0.0016418337,0.0015065902,0.011167302,0.00008795577],"about_ca_topic_score_codex":0.004655227,"about_ca_topic_score_gemma":0.009651348,"teacher_disagreement_score":0.00868684,"about_ca_system_score_codex":0.0013178498,"about_ca_system_score_gemma":0.0019635519,"threshold_uncertainty_score":0.045940936},"labels":[],"label_agreement":null},{"id":"W3097847123","doi":"10.1109/modre51215.2020.00016","title":"A Neural Network Based Approach to Domain Modelling Relationships and Patterns Recognition","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Domain (mathematical analysis); Domain analysis; Domain model; Artificial intelligence; Sentence; Context (archaeology); Artificial neural network; Domain engineering; Exploit; Feature-oriented domain analysis; Domain knowledge; Natural language processing; Machine learning; Software; Software development; Component-based software engineering","score_opus":0.10822412897707545,"score_gpt":0.24515943019828998,"score_spread":0.13693530122121453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3097847123","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012112262,0.00024644294,0.9823873,0.0003529097,0.00006375674,0.00013750026,0.00026197688,0.0015650891,0.0028728375],"genre_scores_gemma":[0.22198193,0.00037194818,0.77021337,0.00028786517,0.00007074772,0.0003328489,0.0010217573,0.00012086823,0.0055986852],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99932456,0.00016562958,0.000054270746,0.00023906108,0.00015938636,0.000057025863],"domain_scores_gemma":[0.99917066,0.0003859358,0.00010109956,0.000100297584,0.00020729787,0.00003476348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008173481,0.00087567884,0.00043730874,0.0014886878,0.0005378174,0.0012115383,0.001362086,0.001254849,0.0026737433],"category_scores_gemma":[0.0031837742,0.00038808555,0.0007484019,0.0018627837,0.0006259269,0.0021374442,0.00091209554,0.0016581236,0.0007933162],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021982823,0.00031455382,0.0032354489,0.00033153905,0.0001643695,0.00057102746,0.00077280233,0.12360287,0.036389705,0.025612954,0.007440421,0.8013444],"study_design_scores_gemma":[0.000007850825,0.00004595517,0.000737953,0.0000255751,0.0000289779,0.00010759205,0.000074610325,0.97296494,0.0060319807,0.016494436,0.0034621102,0.00001804654],"about_ca_topic_score_codex":0.0063168635,"about_ca_topic_score_gemma":0.010422582,"teacher_disagreement_score":0.0063168635,"about_ca_system_score_codex":0.0009079419,"about_ca_system_score_gemma":0.00085963705,"threshold_uncertainty_score":0.012560189},"labels":[],"label_agreement":null},{"id":"W3098139023","doi":"10.1145/3368089.3417922","title":"LibComp: an IntelliJ plugin for comparing Java libraries","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Plug-in; Computer science; Java; Dependency (UML); World Wide Web; Process (computing); Software engineering; Domain (mathematical analysis); Metric (unit); Software; Programming language; Engineering","score_opus":0.09248894072043512,"score_gpt":0.2896499004870143,"score_spread":0.19716095976657916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098139023","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0581275,0.0010557786,0.33050486,0.00037528574,0.00039655314,0.0009929503,0.011426165,0.57811755,0.019003436],"genre_scores_gemma":[0.32453424,0.0006922981,0.5347932,0.00069864537,0.00016119318,0.0018188099,0.035030294,0.08882802,0.013443352],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.995685,0.00085710734,0.0006042991,0.0005972682,0.0019838302,0.0002725551],"domain_scores_gemma":[0.98969895,0.0056064264,0.0010933194,0.0017064605,0.0015083652,0.0003865132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00512101,0.002073788,0.0008199825,0.0071460763,0.00078937534,0.0028922886,0.0024752403,0.001095237,0.010801643],"category_scores_gemma":[0.01873333,0.0013617753,0.0010587884,0.0027522054,0.0008547625,0.0052120197,0.003966036,0.0014723942,0.0037832842],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036558488,0.000997638,0.03625487,0.00320383,0.0005159006,0.0008667918,0.0026573655,0.0060312534,0.03555079,0.021047395,0.22264516,0.66657317],"study_design_scores_gemma":[0.0011920183,0.0021284418,0.119832024,0.0013185112,0.00038957834,0.0029570465,0.0017309285,0.265151,0.13206421,0.032175776,0.4397088,0.0013516975],"about_ca_topic_score_codex":0.0034748788,"about_ca_topic_score_gemma":0.0062975246,"teacher_disagreement_score":0.010801643,"about_ca_system_score_codex":0.00082910596,"about_ca_system_score_gemma":0.0015155828,"threshold_uncertainty_score":0.036135077},"labels":[],"label_agreement":null},{"id":"W3099529967","doi":"10.1109/esem.2019.8870177","title":"On the Impact of Refactoring on the Relationship between Quality Attributes and Design Metrics","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":76,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Quality (philosophy); Software engineering; Programming language; Software","score_opus":0.3890357204351267,"score_gpt":0.2934785192488043,"score_spread":0.09555720118632238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099529967","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9731169,0.0046623684,0.016295178,0.0007313571,0.000049590628,0.000046141755,0.0017888893,0.0002903594,0.0030192935],"genre_scores_gemma":[0.9934036,0.0005185387,0.004142455,0.00005730487,0.00003788173,0.000027952332,0.0014991385,0.00007596493,0.00023714523],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98481154,0.005157664,0.001327686,0.0039661755,0.0041648033,0.0005721506],"domain_scores_gemma":[0.4536672,0.49239323,0.030707438,0.009943803,0.011643302,0.0016449844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01386142,0.0008639202,0.00053154945,0.006186361,0.00050024525,0.0021899804,0.0006866755,0.00131458,0.002060145],"category_scores_gemma":[0.17716311,0.00045834566,0.0010849129,0.005972549,0.0013293434,0.003847084,0.0013909704,0.0023000096,0.0007675538],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023781495,0.00013632998,0.93859667,0.000521581,0.000472427,0.00021005777,0.0009710566,0.0054100803,0.0020058397,0.0005941692,0.0008764368,0.049967464],"study_design_scores_gemma":[0.000009119431,0.0002465481,0.97197163,0.0001547503,0.00026323163,0.00033335804,0.00047631102,0.022651142,0.0013439538,0.0012550668,0.001254551,0.000040279163],"about_ca_topic_score_codex":0.0036104391,"about_ca_topic_score_gemma":0.00400865,"teacher_disagreement_score":0.01386142,"about_ca_system_score_codex":0.00070670753,"about_ca_system_score_gemma":0.0007793189,"threshold_uncertainty_score":0.07330704},"labels":[],"label_agreement":null},{"id":"W3101009354","doi":"10.1109/scam51674.2020.00020","title":"Annotation practices in Android apps","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Canada Research Chairs","keywords":"Android (operating system); Java; Computer science; Annotation; Java Programming Language; World Wide Web; Android app; Empirical research; Mobile device; Application programming interface; Operating system; Artificial intelligence","score_opus":0.05747301230433355,"score_gpt":0.31700793567798646,"score_spread":0.2595349233736529,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3101009354","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9126385,0.0059334403,0.039352436,0.00268881,0.00020286422,0.00057970866,0.00091280416,0.0035137055,0.034177754],"genre_scores_gemma":[0.96125543,0.0022288247,0.026034275,0.0007103829,0.00009489325,0.0004827958,0.0009807395,0.0011402289,0.007072376],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97286093,0.009574669,0.0032245493,0.0041247923,0.009352151,0.0008629584],"domain_scores_gemma":[0.81514347,0.1230931,0.018077081,0.017211001,0.02473297,0.0017424444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0144273965,0.0007039339,0.0005166237,0.0049728462,0.0024357196,0.004263428,0.0012757472,0.0013067612,0.001399713],"category_scores_gemma":[0.103036545,0.0009203421,0.00067350524,0.0030446404,0.002162643,0.007158592,0.004093519,0.0015620489,0.0009370565],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006225975,0.0002590803,0.2856454,0.0026247683,0.00014120206,0.0025907939,0.21892461,0.0010407384,0.022328,0.0071972543,0.008192,0.45043355],"study_design_scores_gemma":[0.00007704842,0.000851303,0.5202575,0.0050580017,0.00055399153,0.007872591,0.100794,0.011115741,0.023021547,0.013548047,0.3159699,0.0008803611],"about_ca_topic_score_codex":0.008278685,"about_ca_topic_score_gemma":0.01126695,"teacher_disagreement_score":0.0144273965,"about_ca_system_score_codex":0.0015714285,"about_ca_system_score_gemma":0.0020694062,"threshold_uncertainty_score":0.07630026},"labels":[],"label_agreement":null},{"id":"W3102320290","doi":"10.22215/etd/2020-14169","title":"Studying the Change History of Code Snippets on Stack Overflow and GitHub","year":2020,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Code (set theory); Reuse; Software evolution; Code reuse; World Wide Web; Source code; Software; Stack (abstract data type); Information retrieval; Code review; Software engineering; Data science; Database; Software development; Programming language; Static program analysis; Software construction; Engineering","score_opus":0.07927431558380178,"score_gpt":0.28968422317857356,"score_spread":0.2104099075947718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3102320290","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97503555,0.00069128396,0.004495384,0.000506171,0.0001391908,0.00008528694,0.006152017,0.0029610912,0.009933943],"genre_scores_gemma":[0.956938,0.00056380266,0.012531051,0.00023536438,0.000081987324,0.00015397243,0.015397083,0.0028267584,0.011271866],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9977239,0.00024709164,0.00009412859,0.00041902586,0.0013129152,0.00020291793],"domain_scores_gemma":[0.9748761,0.011605801,0.004127505,0.0023389335,0.0061476156,0.0009041267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001310451,0.0003532806,0.00030546015,0.0053327642,0.00083252793,0.0012118326,0.0005858743,0.0007549069,0.0024947168],"category_scores_gemma":[0.025302447,0.00037659015,0.00035218973,0.0055401647,0.0008308273,0.0024918849,0.0011581231,0.0012847682,0.0014878184],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011208786,0.0006007423,0.4849542,0.0008504305,0.0002423083,0.0021543733,0.013057002,0.006622181,0.03650952,0.006771887,0.050719265,0.39639717],"study_design_scores_gemma":[0.000041039664,0.0003803841,0.8830809,0.00027676384,0.00013122374,0.0011609953,0.0037316652,0.034668818,0.022105806,0.0036236998,0.050645865,0.00015280212],"about_ca_topic_score_codex":0.0135091385,"about_ca_topic_score_gemma":0.023801306,"teacher_disagreement_score":0.0135091385,"about_ca_system_score_codex":0.0007909879,"about_ca_system_score_gemma":0.0007572763,"threshold_uncertainty_score":0.026861012},"labels":[],"label_agreement":null},{"id":"W3102671910","doi":"","title":"Ammonia: An Approach for Deriving Project Specific Bug Patterns","year":2020,"lang":"en","type":"article","venue":"Institutional Repositories DataBase (IRDB)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software bug; Java; Software engineering; Commit; Security bug; Software development; Software; Programming language; Database; Operating system; Software security assurance","score_opus":0.06341276235819418,"score_gpt":0.2918402594894768,"score_spread":0.2284274971312826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3102671910","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029913355,0.00016045842,0.9195612,0.0002647731,0.00007632449,0.0009386027,0.003807656,0.041952778,0.003324803],"genre_scores_gemma":[0.060426984,0.00012246726,0.9294039,0.00011166871,0.000029025758,0.0007237379,0.004923434,0.001586708,0.0026720613],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99504304,0.0009488976,0.0006331531,0.0012697574,0.0018808934,0.00022413161],"domain_scores_gemma":[0.9876282,0.0044540428,0.002440932,0.0024358798,0.0027350558,0.000305799],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037119777,0.0017742806,0.0009642407,0.008057771,0.0008203791,0.0021485097,0.0020577111,0.0011498815,0.0032636204],"category_scores_gemma":[0.017643984,0.0011517243,0.0017963963,0.0042206394,0.00071604847,0.0025952947,0.0027467324,0.001273637,0.0018993275],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045687836,0.00043332638,0.075636916,0.0015646422,0.00042657572,0.0013443542,0.0034712444,0.015784098,0.0298826,0.010254893,0.023498839,0.8372456],"study_design_scores_gemma":[0.0003247117,0.0010458829,0.06149839,0.0007721207,0.00074611604,0.004356423,0.0025677318,0.68676114,0.0640957,0.037243925,0.14010295,0.00048489988],"about_ca_topic_score_codex":0.0051306877,"about_ca_topic_score_gemma":0.008969431,"teacher_disagreement_score":0.008057771,"about_ca_system_score_codex":0.0006998719,"about_ca_system_score_gemma":0.0027081126,"threshold_uncertainty_score":0.019631088},"labels":[],"label_agreement":null},{"id":"W3102738773","doi":"10.22215/etd/2020-14190","title":"Analysis and Maintainability of Complex Industry Test Code Using Clone Detection","year":2020,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Maintainability; Computer science; Source code; Programming language; Code refactoring; Java; Software maintenance; Code (set theory); Software engineering; Software; Software system; Set (abstract data type); Biology; Genetics","score_opus":0.03408750687434436,"score_gpt":0.3173681645261578,"score_spread":0.28328065765181343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3102738773","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89774555,0.00034258785,0.09839274,0.00005871443,0.000011824528,0.00006907992,0.00023484003,0.0015639877,0.0015807651],"genre_scores_gemma":[0.9523717,0.000111693196,0.045381203,0.000023578603,0.000008673273,0.000041759235,0.0005120361,0.00017375968,0.0013756066],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99876344,0.00018784514,0.000070139606,0.0002535202,0.0006258472,0.00009924738],"domain_scores_gemma":[0.9882923,0.0062386435,0.0020998623,0.0014126657,0.0018274924,0.0001289919],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009995901,0.00035617908,0.00023404654,0.0025992168,0.00024052543,0.000657696,0.00052083965,0.0005678448,0.00096357765],"category_scores_gemma":[0.010671241,0.00022381506,0.0004122792,0.0012105838,0.0005053356,0.00090982544,0.00035292548,0.00036355716,0.00021871326],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006704509,0.0004304721,0.18399893,0.00036740117,0.00017463783,0.00055886497,0.0012223742,0.049714852,0.26301262,0.004370199,0.0014368707,0.49404243],"study_design_scores_gemma":[0.00007364041,0.001154537,0.22968453,0.000070906885,0.00023886157,0.0009477872,0.00016671536,0.558791,0.20211233,0.0035621491,0.0031365394,0.0000611487],"about_ca_topic_score_codex":0.0042829886,"about_ca_topic_score_gemma":0.0043891887,"teacher_disagreement_score":0.0042829886,"about_ca_system_score_codex":0.0004915452,"about_ca_system_score_gemma":0.0005496735,"threshold_uncertainty_score":0.008516073},"labels":[],"label_agreement":null},{"id":"W3104993042","doi":"10.1145/3368089.3417927","title":"JITO: a tool for just-in-time defect identification and localization","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software bug; Task (project management); Identification (biology); Quality assurance; Computer science; Software; Software quality; Quality (philosophy); Software quality assurance; Software engineering; Software quality analyst; Software development; Programming language; Engineering; Systems engineering; Operations management","score_opus":0.022223405175218397,"score_gpt":0.272385967843636,"score_spread":0.25016256266841763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3104993042","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025529386,0.00043136568,0.30973417,0.00013107945,0.00019778725,0.00020181139,0.0042192424,0.6799062,0.002625348],"genre_scores_gemma":[0.08759163,0.0015776347,0.6899092,0.0012437749,0.0003014389,0.0018483369,0.045418143,0.15309256,0.01901736],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99771404,0.00024482055,0.00027559692,0.00050839334,0.0010375835,0.00021961881],"domain_scores_gemma":[0.9951297,0.002292636,0.0005604054,0.0011381581,0.00060066726,0.0002784503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025378608,0.0035396803,0.0018314011,0.0045685396,0.0010359727,0.0025477805,0.003986524,0.0025956922,0.02534748],"category_scores_gemma":[0.00799703,0.0026650098,0.0025561252,0.0017870204,0.00097331894,0.0047761495,0.0045394376,0.0031133238,0.015004264],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018587855,0.0006123708,0.011062772,0.00415916,0.00074110646,0.0019935057,0.0014253929,0.013655978,0.064802,0.012510738,0.43184346,0.45533478],"study_design_scores_gemma":[0.0011869833,0.00083571393,0.014700513,0.0010131388,0.0005670825,0.0042125825,0.0004231581,0.28244734,0.118833564,0.046598934,0.52805287,0.0011281235],"about_ca_topic_score_codex":0.0019005906,"about_ca_topic_score_gemma":0.0025885184,"teacher_disagreement_score":0.02534748,"about_ca_system_score_codex":0.00049474585,"about_ca_system_score_gemma":0.0015713598,"threshold_uncertainty_score":0.08479571},"labels":[],"label_agreement":null},{"id":"W3105904721","doi":"","title":"A Digital Game Maturity Model (DGMM)","year":2016,"lang":"en","type":"article","venue":"Scholarship@Western (Western University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University; Western University","funders":"","keywords":"Maturity (psychological); Video game development; Capability Maturity Model; Popularity; Computer science; Game Developer; Knowledge management; Quality (philosophy); Game design; Software; Process management; Multimedia; Engineering","score_opus":0.07978850938261935,"score_gpt":0.2992522991797711,"score_spread":0.21946378979715175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3105904721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20718263,0.0019767846,0.5858344,0.007706913,0.00020866952,0.005828462,0.0020671075,0.0011731734,0.18802176],"genre_scores_gemma":[0.6544769,0.0010148346,0.33271554,0.00044074588,0.000018658973,0.002982526,0.0012345125,0.000047423906,0.0070689493],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9963414,0.0013067533,0.00041729337,0.0003232073,0.0013395881,0.00027168964],"domain_scores_gemma":[0.9945379,0.0018630882,0.0007756145,0.00032351605,0.0020629785,0.00043690728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005243364,0.00073185517,0.00028867606,0.0033742609,0.0008224612,0.0030546873,0.0013548383,0.0016411274,0.0018167538],"category_scores_gemma":[0.013876119,0.0003120905,0.00082570635,0.001977617,0.00079177856,0.0043634283,0.0021928854,0.0017713769,0.0008781646],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001552449,0.00084828434,0.06493501,0.0011285493,0.00010763718,0.0004844925,0.0089290505,0.01936013,0.008515961,0.45326874,0.011617236,0.4306498],"study_design_scores_gemma":[0.00016921439,0.0020100542,0.07717951,0.0028096412,0.00026158767,0.0029249573,0.013034247,0.22229366,0.010784822,0.3496552,0.31855482,0.00032238985],"about_ca_topic_score_codex":0.0038910839,"about_ca_topic_score_gemma":0.004236217,"teacher_disagreement_score":0.005243364,"about_ca_system_score_codex":0.00434023,"about_ca_system_score_gemma":0.0043816147,"threshold_uncertainty_score":0.031490743},"labels":[],"label_agreement":null},{"id":"W3106278391","doi":"10.1109/issrew.2016.43","title":"On Automatic Detection of Performance Bugs","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software bug; Operating system; Software","score_opus":0.010324885819473264,"score_gpt":0.23211643080949043,"score_spread":0.22179154499001716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3106278391","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51900065,0.0028154256,0.44442472,0.0011099599,0.000284787,0.0004167678,0.0036803442,0.023581684,0.004685717],"genre_scores_gemma":[0.8724393,0.00045494913,0.12050867,0.00020322826,0.00009276967,0.00010705773,0.0038604387,0.00023532414,0.0020982116],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99455243,0.0012058985,0.00048168705,0.0013740883,0.0020918641,0.000293985],"domain_scores_gemma":[0.96290576,0.020370228,0.0068223737,0.0025064358,0.006802571,0.00059265865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003159309,0.0014477645,0.00076674606,0.006065819,0.00063050265,0.001292655,0.0015049637,0.0015571172,0.0011470425],"category_scores_gemma":[0.028992621,0.0004613267,0.000695827,0.0023221583,0.00051554793,0.0022850283,0.0011739531,0.0010752308,0.0014005549],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005357059,0.00080775085,0.31488037,0.00075408246,0.00020967622,0.0007185257,0.00046592273,0.043111753,0.015842656,0.0017910382,0.015717715,0.60516477],"study_design_scores_gemma":[0.00004888092,0.0003630151,0.0587319,0.00013493597,0.00010290088,0.0011615198,0.00018780479,0.9174497,0.014678071,0.0034649768,0.0036081844,0.00006801514],"about_ca_topic_score_codex":0.009182759,"about_ca_topic_score_gemma":0.010194744,"teacher_disagreement_score":0.009182759,"about_ca_system_score_codex":0.00084876426,"about_ca_system_score_gemma":0.0016276669,"threshold_uncertainty_score":0.018258631},"labels":[],"label_agreement":null},{"id":"W3108206695","doi":"10.5753/cbie.sbie.2020.1633","title":"Detection of Programming Plagiarism in Computing Education: A Systematic Mapping Study","year":2020,"lang":"en","type":"article","venue":"Anais do XXXI Simpósio Brasileiro de Informática na Educação (SBIE 2020)","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Coordenação de Aperfeiçoamento de Pessoal de Nível Superior; Canadian Bureau for International Education","keywords":"Computer science; Domain (mathematical analysis); Compromise; Software engineering; Data science; Mathematics","score_opus":0.0199727210368318,"score_gpt":0.2836641455116207,"score_spread":0.26369142447478894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3108206695","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96108145,0.0034651612,0.027542032,0.0004042747,0.000023771328,0.0020865812,0.0010437956,0.0001798175,0.004173242],"genre_scores_gemma":[0.9697954,0.0013854571,0.025344418,0.00012644268,0.000011108989,0.0011789278,0.0010814881,0.000040613028,0.0010361059],"study_design_codex":"observational","study_design_gemma":"systematic_review","domain_scores_codex":[0.982842,0.008536362,0.002537012,0.0022684445,0.0032004435,0.00061564497],"domain_scores_gemma":[0.8770745,0.076950625,0.01594616,0.00870519,0.020090126,0.0012334052],"candidate_categories":["metaresearch","research_integrity"],"consensus_categories":[],"category_scores_codex":[0.016786253,0.0003840004,0.00066448754,0.020413682,0.0020500587,0.0015017578,0.0008904884,0.00078838307,0.001327526],"category_scores_gemma":[0.07841977,0.0003864668,0.0006820243,0.009222671,0.0021193195,0.00259059,0.0027844936,0.000625722,0.00037080015],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026594225,0.0005600953,0.4559268,0.0076979767,0.00027797246,0.0009302745,0.14270464,0.0004941207,0.0076940428,0.0047340724,0.002704835,0.37600926],"study_design_scores_gemma":[0.00009865979,0.001335526,0.75778174,0.0053105364,0.00073845393,0.003471333,0.15653378,0.0048083817,0.017746782,0.007874893,0.04410157,0.00019833392],"about_ca_topic_score_codex":0.003500781,"about_ca_topic_score_gemma":0.0077004367,"teacher_disagreement_score":0.9992116,"about_ca_system_score_codex":0.0014912232,"about_ca_system_score_gemma":0.0060338853,"threshold_uncertainty_score":0.08877522},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"systematic_review","genre":"review","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":["research_integrity"],"domain":null,"study_design":"systematic_review","genre":"review","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"split"},{"id":"W3108760172","doi":"10.1145/3428249","title":"Designing types for R, empirically","year":2020,"lang":"en","type":"article","venue":"Proceedings of the ACM on Programming Languages","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Office of Naval Research; Office of Naval Research Global; Natural Sciences and Engineering Research Council of Canada; European Commission; National Science Foundation","keywords":"Computer science; Compiler; Data type; Programming language; Code (set theory); Variety (cybernetics); Source code; Overhead (engineering); Process (computing); Artificial intelligence","score_opus":0.03649795393744953,"score_gpt":0.30183421922706444,"score_spread":0.2653362652896149,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3108760172","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014001296,0.0011437233,0.94044745,0.010683161,0.00071754516,0.00046339555,0.002092279,0.003828508,0.026622681],"genre_scores_gemma":[0.10899401,0.0006957844,0.87320626,0.0030790837,0.00037379822,0.0016639465,0.0023095089,0.0046669026,0.0050106887],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.8526236,0.11310717,0.007616818,0.011274343,0.013832708,0.0015454048],"domain_scores_gemma":[0.6486421,0.24218497,0.01283214,0.07141099,0.022820728,0.0021090158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.09804514,0.0015059916,0.001449562,0.0047644023,0.0038882273,0.012925783,0.004053406,0.00313942,0.016963985],"category_scores_gemma":[0.45936185,0.002431911,0.0026025777,0.007732223,0.01005544,0.021773214,0.0073275636,0.006506177,0.008561156],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004552044,0.00010743807,0.019796925,0.0014553377,0.00028212913,0.00023952391,0.006129579,0.008227619,0.0016745786,0.72589564,0.05791627,0.17781976],"study_design_scores_gemma":[0.00019660914,0.00017229962,0.0045581944,0.0011359121,0.00014489323,0.00073984923,0.0020257134,0.03814447,0.0037990029,0.73686343,0.21207458,0.00014503564],"about_ca_topic_score_codex":0.0033623385,"about_ca_topic_score_gemma":0.004937627,"teacher_disagreement_score":0.09804514,"about_ca_system_score_codex":0.0033189827,"about_ca_system_score_gemma":0.00778722,"threshold_uncertainty_score":0.51851845},"labels":[],"label_agreement":null},{"id":"W3109504448","doi":"10.1016/j.asoc.2020.106908","title":"WhoReview: A multi-objective search-based approach for code reviewers recommendation in modern code review","year":2020,"lang":"en","type":"article","venue":"Applied Soft Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; École de Technologie Supérieure","funders":"Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada","keywords":"Code review; Computer science; Context (archaeology); Code (set theory); Software development; Set (abstract data type); Software; Workload; Software quality; Source code; Software evolution; Software engineering; Quality (philosophy); Empirical research; Software construction; Programming language","score_opus":0.07410835405611745,"score_gpt":0.3230320341263064,"score_spread":0.24892368007018895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109504448","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036683843,0.017918842,0.8972368,0.0049067973,0.001943304,0.0051769004,0.0086091235,0.016313192,0.011211187],"genre_scores_gemma":[0.1843356,0.0027985626,0.7890131,0.0016410107,0.0013206722,0.0017540853,0.0068241977,0.00083505915,0.011477745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97910917,0.010106132,0.002077152,0.0019832423,0.0061842804,0.000540094],"domain_scores_gemma":[0.9442407,0.036287203,0.0035305417,0.0033976356,0.010947299,0.0015965806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017438522,0.002829801,0.004831679,0.024815693,0.002048118,0.0046834005,0.0039330474,0.0041132937,0.009045043],"category_scores_gemma":[0.058337808,0.0011369826,0.002835602,0.013827657,0.00086958904,0.0043615373,0.0033393358,0.0017391813,0.0030716835],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016007397,0.0013107333,0.012843135,0.005789023,0.0029460674,0.00040262475,0.00067587197,0.029916411,0.007288,0.0068596536,0.116399564,0.8139682],"study_design_scores_gemma":[0.0015486847,0.0019520457,0.008956096,0.0012896575,0.0028322036,0.0010799397,0.0008220414,0.8828286,0.010154285,0.031338338,0.056774832,0.00042324592],"about_ca_topic_score_codex":0.007146368,"about_ca_topic_score_gemma":0.02141778,"teacher_disagreement_score":0.024815693,"about_ca_system_score_codex":0.0014308023,"about_ca_system_score_gemma":0.007161002,"threshold_uncertainty_score":0.09222478},"labels":[],"label_agreement":null},{"id":"W3109776429","doi":"10.1145/3368089.3417057","title":"Establishing key performance indicators for measuring software-development processes at a large organization","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Deutsche Forschungsgemeinschaft","keywords":"Process management; Key (lock); Computer science; Performance indicator; Comparability; Software development; Software; Quality (philosophy); Knowledge management; Engineering management; Risk analysis (engineering); Engineering; Business; Computer security","score_opus":0.024633593611168597,"score_gpt":0.22268670594871995,"score_spread":0.19805311233755135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109776429","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22830708,0.00068799726,0.73425937,0.001603593,0.00026878662,0.0025449067,0.0024043988,0.0028960158,0.027027795],"genre_scores_gemma":[0.50078946,0.00042405914,0.49274573,0.00012033655,0.000081294434,0.0021297294,0.0019509511,0.00038300978,0.0013754645],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.94546396,0.025556536,0.007927781,0.003142481,0.01591869,0.0019906105],"domain_scores_gemma":[0.84600145,0.064650625,0.029473687,0.015404178,0.040049862,0.0044201817],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.048015304,0.0018424892,0.0008514155,0.011626191,0.0018614171,0.006405627,0.0013643034,0.0013305431,0.0022049136],"category_scores_gemma":[0.11835868,0.00048617338,0.0008256342,0.01121065,0.0015508396,0.008801615,0.0037254284,0.0021560038,0.0010056233],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007053792,0.0020540266,0.25145322,0.0025587932,0.00026553473,0.00022784655,0.011825665,0.023656724,0.032470368,0.055415884,0.009077894,0.6102887],"study_design_scores_gemma":[0.00023780362,0.006201071,0.43940744,0.0019477017,0.00056417176,0.0005928945,0.019990094,0.1599186,0.20151618,0.056034755,0.11258755,0.0010017465],"about_ca_topic_score_codex":0.00262297,"about_ca_topic_score_gemma":0.0026562917,"teacher_disagreement_score":0.048015304,"about_ca_system_score_codex":0.0038274534,"about_ca_system_score_gemma":0.0049468265,"threshold_uncertainty_score":0.2539323},"labels":[],"label_agreement":null},{"id":"W3110861511","doi":"10.18608/jla.2020.73.10","title":"Are You Being Rhetorical? A Description of Rhetorical Move Annotation Tools and Open Corpus of Sample Machine-Annotated Rhetorical Moves","year":2020,"lang":"en","type":"article","venue":"Journal of Learning Analytics","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"University of Technology Sydney","keywords":"Rhetorical question; Computer science; Formative assessment; Learning analytics; Annotation; Analytics; Field (mathematics); Sample (material); Data science; Artificial intelligence; Natural language processing; Linguistics; Mathematics education; Psychology","score_opus":0.09093260998804008,"score_gpt":0.3158404757066359,"score_spread":0.22490786571859583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3110861511","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25950596,0.009735046,0.26045084,0.005319886,0.0012992589,0.003215083,0.32787314,0.012760897,0.11983984],"genre_scores_gemma":[0.2161145,0.0036581398,0.4411388,0.00086836796,0.0003183981,0.009009459,0.28834036,0.006370106,0.034181934],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9961824,0.0011133811,0.000619911,0.00080858736,0.0011212652,0.00015441417],"domain_scores_gemma":[0.97484165,0.017538723,0.0013498962,0.002641148,0.0030176984,0.00061082817],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0036073509,0.00066494365,0.00045168208,0.01187716,0.0026149773,0.003633079,0.0009853382,0.0011720848,0.008881205],"category_scores_gemma":[0.020946328,0.00046607284,0.00040894732,0.015692074,0.0018862807,0.0042789136,0.002459894,0.0017543973,0.007178455],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005272422,0.00035241112,0.022942374,0.0079892725,0.000057002435,0.002925575,0.07711463,0.0020733972,0.044293806,0.04647615,0.20755471,0.58769345],"study_design_scores_gemma":[0.000028365956,0.00005388665,0.021300659,0.0012273203,0.00002314568,0.0017068337,0.009932326,0.0027523253,0.010854741,0.009623642,0.94240016,0.00009660815],"about_ca_topic_score_codex":0.002000397,"about_ca_topic_score_gemma":0.0057947277,"teacher_disagreement_score":0.99639267,"about_ca_system_score_codex":0.0011701917,"about_ca_system_score_gemma":0.0023969098,"threshold_uncertainty_score":0.02971059},"labels":[],"label_agreement":null},{"id":"W3111275476","doi":"10.1109/tse.2021.3078384","title":"A Comparison of Natural Language Understanding Platforms for Chatbots in Software Engineering","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":106,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Chatbot; Computer science; IBM; Natural language understanding; Natural language; Task (project management); Software; Data science; Artificial intelligence; Software engineering; World Wide Web; Natural language processing; Programming language; Engineering","score_opus":0.0374386258312366,"score_gpt":0.29901288892110517,"score_spread":0.2615742630898686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111275476","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8653842,0.0070632547,0.06861119,0.0016520058,0.00049672084,0.0012978755,0.006637066,0.033044666,0.01581299],"genre_scores_gemma":[0.8310165,0.0018914195,0.11992458,0.0006229556,0.0001803138,0.0012731258,0.03761604,0.0016653218,0.0058097234],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9900748,0.0045545474,0.0010572454,0.0019354111,0.0019356225,0.000442288],"domain_scores_gemma":[0.9607784,0.029468695,0.0016697082,0.0024487711,0.0039311126,0.0017033222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010499058,0.0015912078,0.0011830279,0.0059866006,0.0011409379,0.002983351,0.0015977412,0.0022929572,0.0021445446],"category_scores_gemma":[0.03641912,0.0005395221,0.0013364104,0.0020345894,0.00090884685,0.00806092,0.003560568,0.002228377,0.002508084],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010056529,0.0030934208,0.052159704,0.007643786,0.0009171618,0.0012050521,0.01802421,0.021246303,0.049362037,0.0088897655,0.044920523,0.7824815],"study_design_scores_gemma":[0.0015009624,0.008492857,0.2226238,0.0019813941,0.00093008636,0.0023382357,0.017687254,0.5328187,0.055389166,0.017167458,0.13794082,0.0011292928],"about_ca_topic_score_codex":0.00555375,"about_ca_topic_score_gemma":0.008249082,"teacher_disagreement_score":0.010499058,"about_ca_system_score_codex":0.0016605905,"about_ca_system_score_gemma":0.0016728784,"threshold_uncertainty_score":0.055525005},"labels":[],"label_agreement":null},{"id":"W3111633890","doi":"10.1109/bigcomp51126.2021.00034","title":"Discovering Business Problems Using Problem Hypotheses: A Goal-Oriented and Machine Learning-Based Approach","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Metis; Computer science; Artificial intelligence; Context (archaeology); Set (abstract data type); Machine learning; Data science; Management science; World Wide Web","score_opus":0.026642883057560907,"score_gpt":0.23774988742307315,"score_spread":0.21110700436551225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111633890","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025807807,0.0005476341,0.96049345,0.005382887,0.00006421634,0.00092817127,0.00054869795,0.0010825642,0.005144608],"genre_scores_gemma":[0.16548507,0.00030997136,0.8310314,0.00059537636,0.000068097055,0.00051693217,0.0010303828,0.000076702585,0.00088603556],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9884985,0.006674824,0.00078014494,0.0014983597,0.0021533282,0.00039488703],"domain_scores_gemma":[0.95172846,0.039567888,0.0027219707,0.002340615,0.0027857644,0.0008551981],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013238529,0.0021518562,0.0015071005,0.0051237973,0.001727763,0.0064835344,0.00414387,0.0026686785,0.0037638329],"category_scores_gemma":[0.035678647,0.0010375802,0.0028888874,0.0026978168,0.005127408,0.008248119,0.0040235957,0.004243612,0.00067687995],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065188884,0.0027714986,0.025253868,0.0032231084,0.00074898254,0.0012277557,0.0047941897,0.19075467,0.006280275,0.28858528,0.013676399,0.46203217],"study_design_scores_gemma":[0.00019521957,0.00037361783,0.0025224593,0.00050326163,0.00020635311,0.000383732,0.002885863,0.6586718,0.0069073965,0.31406108,0.013179936,0.00010927114],"about_ca_topic_score_codex":0.0031838152,"about_ca_topic_score_gemma":0.004857678,"teacher_disagreement_score":0.013238529,"about_ca_system_score_codex":0.0024967538,"about_ca_system_score_gemma":0.0056252535,"threshold_uncertainty_score":0.07001293},"labels":[],"label_agreement":null},{"id":"W3111856980","doi":"10.5121/ijsea.2020.11603","title":"A Data Extraction Algorithm from Open Source Software Project Repositories for Building Duration Estimation Models: Case Study of Github","year":2020,"lang":"en","type":"article","venue":"International Journal of Software Engineering & Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Duration (music); Schedule; Computer science; Estimation; Software; Data extraction; Variable (mathematics); Work (physics); Data mining; Data science; Engineering; Systems engineering; Mathematics","score_opus":0.0748553564177604,"score_gpt":0.3637761955950537,"score_spread":0.2889208391772933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111856980","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1635965,0.0005056139,0.82482016,0.0006501936,0.000055401306,0.00062654907,0.004464392,0.0033589152,0.0019223061],"genre_scores_gemma":[0.23567495,0.00037935167,0.7533501,0.000055091135,0.000027499833,0.0007199503,0.008424677,0.00014388326,0.0012246019],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99810624,0.0005717226,0.00031773697,0.0004488512,0.00045541223,0.00009999591],"domain_scores_gemma":[0.9896081,0.005989944,0.00107482,0.0011558634,0.0019520527,0.00021923942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003468516,0.00084270776,0.0006072738,0.0037619665,0.00071590755,0.0013734741,0.00105925,0.0008902237,0.0008178237],"category_scores_gemma":[0.016126893,0.0004176483,0.0009832499,0.0045732805,0.00032928545,0.0021685192,0.001352038,0.0010491722,0.0008074901],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025136964,0.0006998143,0.083525285,0.000692581,0.00015848,0.00067670294,0.0014644577,0.07403966,0.008369513,0.0067063873,0.008771734,0.81464404],"study_design_scores_gemma":[0.000062228006,0.0002906942,0.035272777,0.00022038801,0.00012592852,0.0006933623,0.0013019963,0.9149073,0.015109298,0.0085084075,0.02340135,0.00010622393],"about_ca_topic_score_codex":0.009077639,"about_ca_topic_score_gemma":0.0132587515,"teacher_disagreement_score":0.009077639,"about_ca_system_score_codex":0.0007928506,"about_ca_system_score_gemma":0.0024004346,"threshold_uncertainty_score":0.018343449},"labels":[],"label_agreement":null},{"id":"W3111871824","doi":"10.1109/smc42975.2020.9283289","title":"Fast Detection of Duplicate Bug Reports using LDA-based Topic Modeling and Classification","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Software regression; Cluster analysis; Latent Dirichlet allocation; Software bug; Topic model; Software; Data mining; Word2vec; Exploit; Software maintenance; Word (group theory); Anomaly detection; Artificial intelligence; Information retrieval; Software system; Software construction; Programming language","score_opus":0.07175247016765941,"score_gpt":0.2792430226897609,"score_spread":0.20749055252210152,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111871824","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20516145,0.0050964504,0.765761,0.0007939075,0.0005622383,0.0004935453,0.00405027,0.015426552,0.0026546062],"genre_scores_gemma":[0.7269439,0.0013394038,0.25119242,0.0002516724,0.00044854413,0.00044658247,0.013539093,0.00044376476,0.005394568],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99629503,0.00090454647,0.00032219396,0.0011162177,0.0010045768,0.00035740994],"domain_scores_gemma":[0.99373007,0.0023111932,0.00060527975,0.0009171189,0.00217178,0.00026468598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002989166,0.002119597,0.0022712476,0.006652176,0.0010504604,0.0019487954,0.002144553,0.0012780572,0.00071320886],"category_scores_gemma":[0.009269791,0.0006016332,0.0023424113,0.0035718884,0.00033235576,0.0023247357,0.0016365718,0.0015174572,0.0022544563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008277498,0.0009201502,0.044682447,0.00042903182,0.00067397323,0.00036924574,0.00074437045,0.029016623,0.017408362,0.0011218331,0.025311016,0.87849516],"study_design_scores_gemma":[0.00004856229,0.00019123688,0.01135161,0.000037488506,0.00022276596,0.00036265666,0.00030339017,0.972969,0.0077086147,0.0021018682,0.0046264436,0.00007642341],"about_ca_topic_score_codex":0.014049095,"about_ca_topic_score_gemma":0.016754072,"teacher_disagreement_score":0.014049095,"about_ca_system_score_codex":0.00065893447,"about_ca_system_score_gemma":0.0014328879,"threshold_uncertainty_score":0.02793461},"labels":[],"label_agreement":null},{"id":"W3114491449","doi":"10.1109/tse.2020.3045914","title":"Revisiting Test Impact Analysis in Continuous Testing From the Perspective of Code Dependencies","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Test Management Approach; Code coverage; Test case; Test (biology); Test harness; Source code; Software quality; Test suite; Regression testing; Code (set theory); Programming language; Reliability engineering; Software development; Software; Software construction; Machine learning; Set (abstract data type); Engineering","score_opus":0.02543541439253849,"score_gpt":0.26515398316104255,"score_spread":0.23971856876850406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3114491449","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6603282,0.0028947373,0.3265298,0.0015981556,0.00012817077,0.00033634907,0.00043527325,0.0019109624,0.0058384305],"genre_scores_gemma":[0.941621,0.00021327096,0.05730719,0.00013139355,0.00007216532,0.00009355217,0.00021571352,0.00014108293,0.00020451743],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9694372,0.015272592,0.0014975171,0.002757887,0.009997563,0.0010371648],"domain_scores_gemma":[0.6016458,0.34249964,0.022385938,0.016338782,0.015303763,0.001826053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01974715,0.0016466517,0.0011040322,0.0063574617,0.0006431938,0.0026745116,0.0030674709,0.0010449657,0.0009992698],"category_scores_gemma":[0.16359049,0.00062189286,0.0009462982,0.004022824,0.0027228596,0.005091826,0.0018592107,0.0028878273,0.00023406083],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008743349,0.0015249122,0.29402488,0.0012881279,0.0006564542,0.0012508607,0.0023570913,0.17318743,0.01970174,0.016200505,0.0030726676,0.485861],"study_design_scores_gemma":[0.00014014651,0.0012286458,0.13822345,0.00047375294,0.00035213178,0.0007906722,0.0011108738,0.81829435,0.0126286,0.023129215,0.003467168,0.00016101378],"about_ca_topic_score_codex":0.0080089215,"about_ca_topic_score_gemma":0.0062582106,"teacher_disagreement_score":0.01974715,"about_ca_system_score_codex":0.0021607925,"about_ca_system_score_gemma":0.0025178327,"threshold_uncertainty_score":0.10443419},"labels":[],"label_agreement":null},{"id":"W3114813405","doi":"10.1007/978-3-030-64694-3_15","title":"How Does Library Migration Impact Software Quality and Comprehension? An Empirical Study","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Software quality; Software engineering; Software development; Software peer review; Software quality analyst; Software quality control; Software; Quality (philosophy); Personal software process; Readability; World Wide Web; Software construction; Operating system; Programming language","score_opus":0.044024959112230644,"score_gpt":0.3310284279706528,"score_spread":0.28700346885842215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3114813405","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9990233,0.00004652473,0.00012618436,0.00006508819,0.0000022649801,0.000010544287,0.000035005654,0.000015364585,0.0006757879],"genre_scores_gemma":[0.9986815,0.000056227225,0.00022391522,0.0000571325,0.000009699246,0.000019832369,0.00009495537,0.000019899151,0.00083697744],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9950058,0.0023952848,0.00042321888,0.0004915754,0.0012888302,0.0003952806],"domain_scores_gemma":[0.801133,0.14592436,0.032034483,0.005240054,0.010094883,0.0055731754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006962177,0.00036749354,0.0006282703,0.001698923,0.0008235552,0.0026993707,0.00080998667,0.0012643478,0.004866854],"category_scores_gemma":[0.07500073,0.00043667914,0.00065704534,0.0019653076,0.0011242148,0.003742932,0.0014836447,0.0015717706,0.0011093629],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015799479,0.0061988737,0.9204987,0.00022537973,0.00014218804,0.00034330433,0.030076995,0.00046505456,0.0024344863,0.0003172068,0.0010267567,0.036691003],"study_design_scores_gemma":[0.00008929509,0.0031399638,0.9703789,0.00007323364,0.00016809002,0.00026452975,0.019528113,0.0019596608,0.0023971703,0.0006266428,0.0013175426,0.000056836983],"about_ca_topic_score_codex":0.0039613764,"about_ca_topic_score_gemma":0.0038017463,"teacher_disagreement_score":0.006962177,"about_ca_system_score_codex":0.0012244111,"about_ca_system_score_gemma":0.0013214357,"threshold_uncertainty_score":0.036819994},"labels":[],"label_agreement":null},{"id":"W3117161066","doi":"10.1145/3412378","title":"An Empirical Study of Developer Discussions in the Gitter Platform","year":2020,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada; Queen's University","funders":"","keywords":"Popularity; Computer science; Thread (computing); World Wide Web; Psychology","score_opus":0.18110658069790245,"score_gpt":0.38313018133306537,"score_spread":0.20202360063516292,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3117161066","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99774206,0.000114833136,0.000918003,0.000070662296,0.0000073011356,0.00010354432,0.00026172926,0.00003606854,0.00074590463],"genre_scores_gemma":[0.9944812,0.0001481357,0.0032201235,0.000069919515,0.000030601732,0.00029034915,0.00085648557,0.00003502408,0.00086819835],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99085265,0.0047889072,0.00068650383,0.0012361266,0.0019391247,0.0004966398],"domain_scores_gemma":[0.8524474,0.11573511,0.016007341,0.0034235085,0.00967649,0.002710197],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009413514,0.00045900588,0.0004693137,0.004232516,0.0014398029,0.0012598946,0.000718105,0.00078785996,0.00081582874],"category_scores_gemma":[0.060946953,0.00036684563,0.00029502827,0.0032079588,0.0011652923,0.0019044417,0.0012309614,0.0010687796,0.00042347502],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006651548,0.0011576838,0.73002625,0.0011201401,0.00012831677,0.0013828072,0.17201348,0.0008263594,0.009540656,0.0008344975,0.0038168244,0.07848784],"study_design_scores_gemma":[0.000046676156,0.00086724386,0.93208617,0.0002336244,0.00005205721,0.00062660803,0.04609037,0.005313574,0.003527413,0.00039960057,0.010672079,0.00008445884],"about_ca_topic_score_codex":0.0031250163,"about_ca_topic_score_gemma":0.006092117,"teacher_disagreement_score":0.009413514,"about_ca_system_score_codex":0.0012036694,"about_ca_system_score_gemma":0.0008734238,"threshold_uncertainty_score":0.049784005},"labels":[],"label_agreement":null},{"id":"W3118008377","doi":"10.1109/qrs-c51114.2020.00070","title":"QMine: A Framework for Mining Quantitative Regular Expressions from System Traces","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Task (project management); Regular expression; Property (philosophy); Data mining; Software system; Software; Cyber-physical system; Programming language; Engineering; Systems engineering","score_opus":0.06514421207624957,"score_gpt":0.31909752651988094,"score_spread":0.25395331444363134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118008377","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010958305,0.0007013475,0.9492862,0.00041909638,0.000051666746,0.0005041822,0.013919432,0.023015777,0.0011438706],"genre_scores_gemma":[0.09727761,0.00053444854,0.86270946,0.00024951992,0.00006553064,0.00091321487,0.035609357,0.0011433025,0.0014974801],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966331,0.00068608625,0.0004468925,0.0008824068,0.0011812531,0.00017031064],"domain_scores_gemma":[0.993012,0.00402949,0.00082751655,0.0011634772,0.0008282003,0.00013927078],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028630567,0.0016623464,0.0010354448,0.004882599,0.0006056178,0.0016937789,0.0026435682,0.0011158946,0.0020086628],"category_scores_gemma":[0.018645605,0.00070771144,0.0021067655,0.0028494948,0.0008784587,0.0027865348,0.0018716253,0.0018106945,0.0011205694],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005864493,0.00078767026,0.032149497,0.0038308732,0.0009063341,0.0013435356,0.0010684105,0.19088723,0.023487527,0.0557503,0.06734523,0.621857],"study_design_scores_gemma":[0.00012385524,0.00024746906,0.0047363047,0.00018903581,0.00011079561,0.0007730765,0.00025723607,0.8749515,0.013416419,0.059623014,0.045476772,0.000094451265],"about_ca_topic_score_codex":0.006115175,"about_ca_topic_score_gemma":0.014884178,"teacher_disagreement_score":0.006115175,"about_ca_system_score_codex":0.00090046634,"about_ca_system_score_gemma":0.0025577706,"threshold_uncertainty_score":0.0151414275},"labels":[],"label_agreement":null},{"id":"W3118322826","doi":"10.1007/s10664-020-09921-9","title":"ID-correspondence: a measure for detecting evolutionary coupling","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Identifier; False positive paradox; Computer science; Similarity (geometry); Measure (data warehouse); Artificial intelligence; Coupling (piping); Data mining; Natural language processing; Machine learning; Programming language; Engineering","score_opus":0.032689246102338025,"score_gpt":0.29433416246353283,"score_spread":0.2616449163611948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118322826","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25922307,0.00084074924,0.7190608,0.00024741283,0.00021937632,0.0002827004,0.0023541895,0.0036121244,0.014159644],"genre_scores_gemma":[0.7627188,0.00023954126,0.23088217,0.00013731503,0.00014871937,0.0002934081,0.0021202324,0.00037777715,0.0030819264],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943705,0.0013344486,0.00047846537,0.0011315387,0.0023800703,0.00030501786],"domain_scores_gemma":[0.9659948,0.019305227,0.00435426,0.0046694903,0.0043750945,0.001301148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004167564,0.0009706504,0.0011070215,0.014222747,0.0014907023,0.0022664238,0.0020207441,0.0019445793,0.0042413697],"category_scores_gemma":[0.038336087,0.00039036246,0.00086960936,0.007499421,0.0014072533,0.004036795,0.0030094679,0.0016616026,0.0013655598],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013314007,0.0008728829,0.27072677,0.0011394902,0.0008424676,0.0007797595,0.0016083423,0.042295944,0.044483438,0.081563294,0.012441875,0.54191434],"study_design_scores_gemma":[0.0001317845,0.0011988839,0.13794948,0.00019299818,0.00042664885,0.0034924648,0.0012287205,0.60587853,0.050838057,0.1777085,0.020591646,0.0003622609],"about_ca_topic_score_codex":0.0011114955,"about_ca_topic_score_gemma":0.0010735727,"teacher_disagreement_score":0.014222747,"about_ca_system_score_codex":0.0009141955,"about_ca_system_score_gemma":0.0011245318,"threshold_uncertainty_score":0.022040486},"labels":[],"label_agreement":null},{"id":"W3118363209","doi":"10.1007/s40614-020-00273-9","title":"Machine Learning for Supplementing Behavioral Assessment","year":2021,"lang":"en","type":"article","venue":"Perspectives on Behavior Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Artificial intelligence; Machine learning; Computer science; Artificial neural network; Function (biology); Experimental data; Statistics; Mathematics","score_opus":0.038319063197590895,"score_gpt":0.3823425613450445,"score_spread":0.34402349814745364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118363209","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.073328584,0.0006639552,0.91448486,0.0015094022,0.00021511932,0.0001657049,0.00051360746,0.0012225067,0.007896391],"genre_scores_gemma":[0.7557746,0.0003642402,0.23895507,0.0002557096,0.00013262939,0.00021035982,0.00035492552,0.00006978814,0.0038825474],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99685115,0.0020545607,0.000115514384,0.00032231753,0.00056994543,0.00008652651],"domain_scores_gemma":[0.97098464,0.022959197,0.0010862434,0.0023113093,0.0023948895,0.00026379398],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044397237,0.00075152004,0.00086530345,0.0014705824,0.0002743945,0.0012260117,0.0009939583,0.0011157485,0.0033552835],"category_scores_gemma":[0.033807784,0.00020471467,0.00036872865,0.0011724221,0.0005575207,0.0021947601,0.00082020985,0.0014508242,0.0007964612],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030111722,0.0010017947,0.029846124,0.00026819992,0.00019706483,0.00009600217,0.00021521508,0.093643695,0.004573264,0.031216366,0.0036889624,0.8349521],"study_design_scores_gemma":[0.000015958216,0.00024233412,0.0067339335,0.00008979986,0.000047353664,0.00006292659,0.00007665286,0.92658126,0.003987801,0.05961203,0.0025150517,0.000034764755],"about_ca_topic_score_codex":0.0028972852,"about_ca_topic_score_gemma":0.0033123202,"teacher_disagreement_score":0.0044397237,"about_ca_system_score_codex":0.0005042302,"about_ca_system_score_gemma":0.0009143265,"threshold_uncertainty_score":0.02347982},"labels":[],"label_agreement":null},{"id":"W3118479550","doi":"10.1007/978-3-030-65796-3_40","title":"A Framework for the Semiotic Quality of User Stories","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Pacific Railway (Canada); Concordia University","funders":"","keywords":"User story; Agile software development; Quality (philosophy); Computer science; Semiotics; User experience design; Context (archaeology); User requirements document; User Research; World Wide Web; Human–computer interaction; Software; Software development; Software engineering; Linguistics","score_opus":0.04557097191535084,"score_gpt":0.30741007594455066,"score_spread":0.26183910402919985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118479550","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012511106,0.0014617098,0.86891484,0.0055934964,0.00022912334,0.00019860464,0.0002708666,0.00070661766,0.11011365],"genre_scores_gemma":[0.6228667,0.0006974661,0.36071593,0.0004127639,0.00024684495,0.0006504702,0.0003959077,0.00054584537,0.013468097],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98900956,0.007482691,0.00064295426,0.00087477267,0.0015198037,0.00047025905],"domain_scores_gemma":[0.97448957,0.01717211,0.0011609893,0.0038300944,0.0024597044,0.00088750146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0108725345,0.000962977,0.00076443003,0.004761123,0.0037753521,0.015227921,0.0028268262,0.0031633691,0.0121315075],"category_scores_gemma":[0.031488508,0.0011371386,0.0011912363,0.0034531346,0.024019863,0.024097746,0.0063415198,0.004294304,0.0013367236],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009821142,0.0000044875956,0.00006998424,0.000028563398,0.0000024729782,0.000031086933,0.003696613,0.0002632645,0.00015791418,0.99083954,0.0005162643,0.0043800557],"study_design_scores_gemma":[0.000012336669,0.000015810818,0.00012091141,0.000073910465,0.000007469734,0.00015634853,0.003039233,0.0049835504,0.0004018343,0.9685875,0.022584924,0.000016223143],"about_ca_topic_score_codex":0.003696455,"about_ca_topic_score_gemma":0.002458842,"teacher_disagreement_score":0.015227921,"about_ca_system_score_codex":0.0044824015,"about_ca_system_score_gemma":0.0023775818,"threshold_uncertainty_score":0.057500124},"labels":[],"label_agreement":null},{"id":"W3118614203","doi":"10.1109/tse.2020.3048335","title":"Accelerating Continuous Integration by Caching Environments and Inferring Dependencies","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; McGill University","funders":"Mitacs","keywords":"Computer science; Acceleration; Service (business); Dependency (UML); Task (project management); Distributed computing; Process (computing); Software; Software engineering; Operating system; Systems engineering","score_opus":0.018451724986903898,"score_gpt":0.2205484944750087,"score_spread":0.20209676948810482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118614203","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48712516,0.002816434,0.36979514,0.001197232,0.00014945882,0.00038634503,0.027522596,0.10310407,0.007903511],"genre_scores_gemma":[0.49232423,0.0007277731,0.42801517,0.00025405336,0.00005253864,0.00027565742,0.07251174,0.0038227926,0.0020161318],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9952226,0.00080126507,0.0004635814,0.0013656899,0.0017858652,0.000361008],"domain_scores_gemma":[0.985,0.0051060785,0.0017118364,0.0056002075,0.0021833773,0.0003984941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029848802,0.0018262967,0.00084282097,0.005907256,0.0008571355,0.002557745,0.0030265686,0.00085392967,0.0010374404],"category_scores_gemma":[0.01995222,0.0017468152,0.0014439244,0.0060351263,0.0010651865,0.004701651,0.0030815701,0.0020390917,0.0013993818],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008130749,0.00062250334,0.26221693,0.0012191227,0.0006241226,0.0008492937,0.0028962488,0.1735685,0.028376367,0.010871323,0.046245087,0.47169745],"study_design_scores_gemma":[0.00008259677,0.00022638176,0.04543383,0.00014756293,0.00031789686,0.0004133298,0.00063339924,0.8684969,0.031470165,0.014163469,0.038451366,0.00016309803],"about_ca_topic_score_codex":0.019033477,"about_ca_topic_score_gemma":0.04826032,"teacher_disagreement_score":0.019033477,"about_ca_system_score_codex":0.0011417118,"about_ca_system_score_gemma":0.002772157,"threshold_uncertainty_score":0.037845433},"labels":[],"label_agreement":null},{"id":"W3119296503","doi":"10.3390/electronics10020179","title":"Empirical Analysis of Rank Aggregation-Based Multi-Filter Feature Selection Methods in Software Defect Prediction","year":2021,"lang":"en","type":"article","venue":"Electronics","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Yayasan UTP; Universiti Teknologi Petronas","keywords":"Rank (graph theory); Filter (signal processing); Computer science; Feature selection; Selection (genetic algorithm); Data mining; Artificial intelligence; Pattern recognition (psychology); Machine learning; Mathematics","score_opus":0.025016988605029693,"score_gpt":0.3487199560671791,"score_spread":0.3237029674621494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119296503","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.394242,0.0041369717,0.5980336,0.0005404877,0.00010156955,0.00012978015,0.00046285993,0.0008201905,0.0015325535],"genre_scores_gemma":[0.94006276,0.0003503448,0.058187753,0.000065278764,0.00006925415,0.00007013093,0.0005612591,0.00003528033,0.00059795234],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967608,0.0011973631,0.0002499767,0.0005270963,0.0010611751,0.00020353879],"domain_scores_gemma":[0.985671,0.009755192,0.0011283277,0.001016651,0.0022131219,0.00021567458],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00788623,0.0009183397,0.0010868546,0.0028007592,0.00049730926,0.0009984495,0.00080445607,0.000858101,0.0006780535],"category_scores_gemma":[0.017153533,0.000213764,0.0009808714,0.0017410077,0.0005088045,0.0015046914,0.00051443215,0.0008960655,0.00021788092],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067040237,0.0005125261,0.0950354,0.00038068584,0.00069160014,0.00019218565,0.00024119967,0.30646107,0.005854884,0.003516948,0.0050320425,0.58141106],"study_design_scores_gemma":[0.000019486186,0.00023781045,0.015071833,0.000027359329,0.00008483712,0.000111626316,0.000055801313,0.9799483,0.0023098784,0.0015009028,0.00060422014,0.000028023313],"about_ca_topic_score_codex":0.004023489,"about_ca_topic_score_gemma":0.0037012484,"teacher_disagreement_score":0.00788623,"about_ca_system_score_codex":0.00061364146,"about_ca_system_score_gemma":0.00087620044,"threshold_uncertainty_score":0.04170686},"labels":[],"label_agreement":null},{"id":"W3119686444","doi":"10.1007/s10664-021-10024-2","title":"A pragmatic approach for hyper-parameter tuning in search-based test case generation","year":2021,"lang":"en","type":"preprint","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Metric (unit); Computer science; Heuristic; Domain (mathematical analysis); Fine-tuning; Class (philosophy); Test case; Parameter space; Performance metric; Mathematical optimization; Machine learning; Artificial intelligence; Mathematics; Statistics; Engineering","score_opus":0.06445494654789755,"score_gpt":0.31047606952561757,"score_spread":0.24602112297772002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119686444","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0034694953,0.000067119894,0.9909554,0.00037937597,0.000036663565,0.00045340558,0.00005442647,0.0018271296,0.0027568927],"genre_scores_gemma":[0.16941792,0.00005001345,0.82634354,0.00064033817,0.00005225139,0.0012079576,0.00015243418,0.0007091593,0.0014264085],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96605676,0.022323484,0.0018038033,0.0018997175,0.007101922,0.0008143376],"domain_scores_gemma":[0.9330552,0.04865988,0.0018684092,0.0092816595,0.006240278,0.0008945048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01881284,0.0016079647,0.0018881515,0.0028395026,0.0015882319,0.004423411,0.004110392,0.004361264,0.011379559],"category_scores_gemma":[0.12222245,0.0017194066,0.0014232418,0.0016390199,0.002640289,0.0040041385,0.0067847827,0.0048427386,0.0027837881],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017060521,0.0011536578,0.0031279074,0.0010608375,0.00034103933,0.00073445635,0.0024626662,0.0882841,0.0521111,0.15341105,0.0129438685,0.68266314],"study_design_scores_gemma":[0.0006419237,0.0003391444,0.0009517655,0.00024970528,0.00015560945,0.00045076502,0.00034151482,0.8399852,0.017809669,0.12770575,0.011216202,0.00015281758],"about_ca_topic_score_codex":0.0011206875,"about_ca_topic_score_gemma":0.0022271548,"teacher_disagreement_score":0.01881284,"about_ca_system_score_codex":0.0012336834,"about_ca_system_score_gemma":0.0031860857,"threshold_uncertainty_score":0.09949297},"labels":[],"label_agreement":null},{"id":"W3119769340","doi":"10.1007/s10664-020-09917-5","title":"An exploratory study on the introduction and removal of different types of technical debt in deep learning frameworks","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Technical debt; Debt; Deep learning; Exploratory research; Quality (philosophy); Field (mathematics)","score_opus":0.018401164916278124,"score_gpt":0.2802219691050121,"score_spread":0.261820804188734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119769340","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9960477,0.00009595014,0.0010377856,0.0001814843,0.000006828295,0.000047683625,0.00004161559,0.000019894236,0.002520988],"genre_scores_gemma":[0.9971117,0.00005574958,0.0015245653,0.00010707793,0.0000069518674,0.000037583985,0.000096541655,0.000018081884,0.0010418771],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99445146,0.0025624607,0.0003225613,0.000516574,0.0013703045,0.0007766506],"domain_scores_gemma":[0.88644075,0.08443769,0.013356247,0.0067346706,0.0052753664,0.003755264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009577751,0.00026312514,0.00033526542,0.0009399501,0.0013848837,0.0020710505,0.0015099309,0.0015533358,0.0027383817],"category_scores_gemma":[0.079654366,0.0003401849,0.00028452635,0.0012093172,0.0015876883,0.0040442115,0.0018151572,0.0037374455,0.00028474536],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039005084,0.016867433,0.55453116,0.0012081633,0.00024266074,0.0020317764,0.082291596,0.0058345203,0.021201082,0.036476538,0.0048302435,0.27058432],"study_design_scores_gemma":[0.00042504008,0.0066607753,0.78640085,0.0008235967,0.00030192963,0.0014041702,0.09247161,0.029159464,0.018194132,0.018531803,0.045395605,0.00023103088],"about_ca_topic_score_codex":0.0029826972,"about_ca_topic_score_gemma":0.005832723,"teacher_disagreement_score":0.009577751,"about_ca_system_score_codex":0.0019435274,"about_ca_system_score_gemma":0.0019391119,"threshold_uncertainty_score":0.050652623},"labels":[],"label_agreement":null},{"id":"W3119800663","doi":"10.1109/tse.2020.3048991","title":"A Method to Assess and Argue for Practical Significance in Software Engineering","year":2018,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Marcus och Amalia Wallenbergs minnesfond","keywords":"Computer science; Bayesian probability; Statistical hypothesis testing; Software; Context (archaeology); Machine learning; Data mining; Empirical research; Statistical model; Data science; Artificial intelligence; Statistics; Mathematics; Programming language","score_opus":0.12360037656386688,"score_gpt":0.27344290841784996,"score_spread":0.14984253185398308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119800663","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013895833,0.00020967658,0.99004185,0.002706288,0.00015656347,0.00025037056,0.00011857714,0.00026315215,0.004863819],"genre_scores_gemma":[0.07728317,0.00035950638,0.9164883,0.001353772,0.0004296626,0.0023205723,0.00016548818,0.0002347127,0.0013648648],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.85776985,0.11059346,0.0046922453,0.008690436,0.017185085,0.0010690587],"domain_scores_gemma":[0.5430692,0.39941728,0.014222015,0.027999554,0.012886705,0.0024051697],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.13563149,0.0027536193,0.0031763073,0.013843565,0.0045617833,0.012617972,0.0059683206,0.007779303,0.015741002],"category_scores_gemma":[0.33579355,0.0020150354,0.0052475017,0.008715441,0.024861801,0.018491196,0.013479577,0.016991897,0.003224778],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005452417,0.00007700809,0.0012394162,0.0002996074,0.00014017684,0.00014048706,0.001497492,0.0028148287,0.00037542076,0.959627,0.0027814726,0.030952547],"study_design_scores_gemma":[0.00004128825,0.00006874065,0.00035126522,0.00022795353,0.00004323474,0.00011210293,0.00022346053,0.015494138,0.00033284933,0.972562,0.010499797,0.000043142438],"about_ca_topic_score_codex":0.0023594955,"about_ca_topic_score_gemma":0.0017942649,"teacher_disagreement_score":0.8643685,"about_ca_system_score_codex":0.0050721164,"about_ca_system_score_gemma":0.00871367,"threshold_uncertainty_score":0.7172964},"labels":[],"label_agreement":null},{"id":"W3120265939","doi":"10.1109/issrew51248.2020.00033","title":"Applying Modular Decomposition in Simulink","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Modular design; Cohesion (chemistry); Cyclomatic complexity; Aerospace; Software engineering; Testability; Decomposition; Software; Model-based testing; Computer architecture; Programming language; Reliability engineering; Test case; Engineering","score_opus":0.024846437638358216,"score_gpt":0.285649911631462,"score_spread":0.2608034739931038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3120265939","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009424576,0.000029533097,0.98656553,0.000029103441,0.000021720505,0.000035044137,0.000024861298,0.0011377796,0.0027318182],"genre_scores_gemma":[0.27092704,0.00015876145,0.726157,0.00004203852,0.000018175713,0.0001367229,0.00014341329,0.00046438217,0.0019524324],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99917275,0.00031825548,0.000053770906,0.00009742849,0.0002916204,0.000066254055],"domain_scores_gemma":[0.99813557,0.0009393078,0.00016419786,0.00036740312,0.00033907825,0.000054441327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001416709,0.00072348077,0.0003331987,0.0006853359,0.00025210858,0.00082898507,0.0005356197,0.00041644607,0.0031013668],"category_scores_gemma":[0.0046443758,0.00036694147,0.0007207906,0.00035740668,0.00049625704,0.00091735105,0.0012168539,0.00087572593,0.0008665262],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015516023,0.000119745986,0.0019747494,0.0002210361,0.00008714822,0.0003258282,0.0004355151,0.7561742,0.033889227,0.09482929,0.0009774711,0.11081073],"study_design_scores_gemma":[0.00003526072,0.00009456501,0.00016625968,0.000034925284,0.000033739263,0.00007062028,0.000029909352,0.95773125,0.014941512,0.018480513,0.008369515,0.00001189739],"about_ca_topic_score_codex":0.0016159882,"about_ca_topic_score_gemma":0.0011286013,"teacher_disagreement_score":0.0031013668,"about_ca_system_score_codex":0.000353509,"about_ca_system_score_gemma":0.0006284012,"threshold_uncertainty_score":0.010375142},"labels":[],"label_agreement":null},{"id":"W3121534148","doi":"10.2139/ssrn.1081429","title":"Business Modeling to Improve Auditor Risk Assessment: An Investigation of Alternative Representations","year":2008,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge; University of Waterloo","funders":"","keywords":"Diagrammatic reasoning; Audit; Presentation (obstetrics); Financial statement; Structuring; Accounting; Representation (politics); Computer science; Statement (logic); Psychology; Knowledge management; Business; Finance; Linguistics; Medicine","score_opus":0.01958356142595293,"score_gpt":0.3009058046629287,"score_spread":0.2813222432369758,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3121534148","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12080541,0.00041582348,0.85296273,0.0036144014,0.000094288196,0.00019864002,0.00042725148,0.00076963124,0.020711707],"genre_scores_gemma":[0.7591841,0.00035927395,0.23774987,0.00019749283,0.000039422663,0.000102302685,0.00027591828,0.00012696176,0.0019646115],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99552804,0.002713762,0.00019350063,0.00028135438,0.0010980292,0.00018539849],"domain_scores_gemma":[0.9789651,0.013542397,0.0015990022,0.0035052248,0.0020863335,0.00030193434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074777487,0.00073011086,0.000741561,0.0017350024,0.000650587,0.005102323,0.0021048307,0.0011221285,0.0035647003],"category_scores_gemma":[0.039134223,0.0003689287,0.0011669802,0.0021902586,0.0007674229,0.005878203,0.0013543301,0.0018459555,0.00042020492],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033562904,0.0005156848,0.0063314904,0.00018841839,0.00020667491,0.00010801611,0.00089092186,0.29568252,0.00097245193,0.54411286,0.003328346,0.14732698],"study_design_scores_gemma":[0.000033202345,0.00006358805,0.0005634024,0.000064766005,0.00006349781,0.00004140323,0.00015949183,0.86079454,0.00053468434,0.1353208,0.0023416306,0.000019089372],"about_ca_topic_score_codex":0.0050480003,"about_ca_topic_score_gemma":0.00504361,"teacher_disagreement_score":0.0074777487,"about_ca_system_score_codex":0.0015643247,"about_ca_system_score_gemma":0.0025544334,"threshold_uncertainty_score":0.03954655},"labels":[],"label_agreement":null},{"id":"W3121596715","doi":"10.1007/s10664-017-9521-5","title":"Do developers update their library dependencies?","year":2017,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":335,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science","keywords":"Reuse; Dependency (UML); Computer science; Workload; Exploit; Software; World Wide Web; Software engineering; Data science; Computer security; Engineering","score_opus":0.03368879006535396,"score_gpt":0.27991691801105534,"score_spread":0.24622812794570137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3121596715","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9548504,0.00075056386,0.005463851,0.005922866,0.00011510686,0.00006150495,0.0012215924,0.0006858766,0.030928278],"genre_scores_gemma":[0.98866165,0.00036631676,0.0021991646,0.00088197464,0.000053067815,0.00003228203,0.0006989458,0.00033027717,0.0067763985],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9903328,0.002700974,0.0008634981,0.0012044355,0.003983513,0.0009148067],"domain_scores_gemma":[0.6781005,0.17249158,0.074234806,0.034428388,0.034123104,0.0066215578],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010384428,0.0004406657,0.0004082184,0.0037461696,0.0013201408,0.0035994942,0.0016320902,0.0019040348,0.0099123735],"category_scores_gemma":[0.21997064,0.0008969265,0.00032431528,0.0033342899,0.0014413485,0.009339746,0.0017561277,0.0022612917,0.0023015453],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028912505,0.00048653272,0.8194011,0.00021580106,0.00011984352,0.0005933955,0.011374808,0.00057392305,0.0014679356,0.0053179683,0.013480783,0.14667885],"study_design_scores_gemma":[0.000105552186,0.00026019613,0.9167011,0.00042440605,0.00034959117,0.0016638163,0.016427245,0.004835694,0.0052137836,0.012238606,0.041661892,0.00011804588],"about_ca_topic_score_codex":0.014999587,"about_ca_topic_score_gemma":0.02827103,"teacher_disagreement_score":0.98961556,"about_ca_system_score_codex":0.0019584303,"about_ca_system_score_gemma":0.0038065747,"threshold_uncertainty_score":0.054918766},"labels":[],"label_agreement":null},{"id":"W3121922233","doi":"10.7287/peerj.preprints.1597","title":"Multi-token code suggestions using statistical language models","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Naturalness; Security token; Computer science; Surprise; Programmer; Code (set theory); Metric (unit); Programming language; Simple (philosophy); Natural language processing; Psychology; Operating system","score_opus":0.1135364606180074,"score_gpt":0.3599472469151834,"score_spread":0.24641078629717603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3121922233","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04116384,0.00015186978,0.9202168,0.00070775184,0.00013928355,0.00009880037,0.0005363196,0.03587763,0.0011076774],"genre_scores_gemma":[0.32075843,0.00009259662,0.67288125,0.00022546746,0.00007382813,0.00012958724,0.0008558543,0.0022297855,0.0027531707],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99757224,0.0012043285,0.00011200056,0.0004894623,0.00052063074,0.00010128014],"domain_scores_gemma":[0.9772379,0.017273841,0.0012600811,0.0017764891,0.0020831747,0.00036846707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034474323,0.0011745951,0.0007214476,0.0011734483,0.00058782444,0.001519067,0.0021725667,0.0010044379,0.0045832093],"category_scores_gemma":[0.02331807,0.00060961314,0.0008971594,0.0006933893,0.0006792027,0.0030845175,0.0011701153,0.0020995021,0.002367064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024870164,0.0005519444,0.015722096,0.0011037878,0.0003074935,0.0013725999,0.0019005489,0.26089752,0.036670636,0.032178864,0.029988939,0.61681855],"study_design_scores_gemma":[0.00003788641,0.00006738034,0.000362945,0.000013007591,0.000018646557,0.00009180424,0.000053650503,0.9794728,0.0065825256,0.010346764,0.0029202886,0.000032276115],"about_ca_topic_score_codex":0.0044920403,"about_ca_topic_score_gemma":0.012730773,"teacher_disagreement_score":0.0045832093,"about_ca_system_score_codex":0.00076886016,"about_ca_system_score_gemma":0.0018655502,"threshold_uncertainty_score":0.018231988},"labels":[],"label_agreement":null},{"id":"W3122365372","doi":"10.7287/peerj.preprints.1132","title":"Error location in Python: where the mutants hide","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Programming language; Python (programming language); Scripting language; Syntax error; Java; Syntax; Abstract syntax; Abstract syntax tree; Static analysis; Compiled language; Compiler; Programming paradigm; High-level programming language; Artificial intelligence; Parsing; Semantics (computer science)","score_opus":0.04883278863831983,"score_gpt":0.30199300757356085,"score_spread":0.253160218935241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122365372","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12329235,0.0010239473,0.7330556,0.0039041783,0.0011767731,0.00018158079,0.0013697963,0.11749586,0.018499933],"genre_scores_gemma":[0.6561125,0.00080402254,0.26744506,0.002882692,0.00017878169,0.0002527125,0.0013644275,0.053131394,0.017828379],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99560803,0.0008030517,0.00044113523,0.000965169,0.0016867693,0.00049587927],"domain_scores_gemma":[0.9865616,0.00372212,0.0016217562,0.0061888555,0.0013876153,0.0005179175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003304476,0.0009864616,0.00093387946,0.0008112071,0.001450057,0.0027149725,0.0025419814,0.0018226748,0.006786226],"category_scores_gemma":[0.023614531,0.0010463418,0.0012734544,0.0007126809,0.0035148854,0.008713583,0.005356509,0.0036119663,0.003970163],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020308767,0.00069148844,0.041456997,0.0019279844,0.00023775133,0.0054719425,0.0075478028,0.02247271,0.14239673,0.24860615,0.056347378,0.47081214],"study_design_scores_gemma":[0.00015858843,0.00042047873,0.009038985,0.0013372628,0.00030776393,0.004087827,0.0016802404,0.096554026,0.33927828,0.27849376,0.26806593,0.0005768648],"about_ca_topic_score_codex":0.0019639865,"about_ca_topic_score_gemma":0.0021611464,"teacher_disagreement_score":0.006786226,"about_ca_system_score_codex":0.0008781575,"about_ca_system_score_gemma":0.0036396505,"threshold_uncertainty_score":0.022702217},"labels":[],"label_agreement":null},{"id":"W3122409631","doi":"10.1145/3324884.3416544","title":"Predicting code context models for software development tasks","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; KPI-driven code analysis; Source code; Leverage (statistics); Code review; Code (set theory); Context (archaeology); Software development; Static program analysis; Software; Eclipse; Codebase; Programming language; Artificial intelligence; Set (abstract data type)","score_opus":0.06418013791197634,"score_gpt":0.2746239992849968,"score_spread":0.21044386137302046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122409631","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82747704,0.0009829863,0.16458274,0.000360133,0.000049568116,0.00014191288,0.0024364123,0.0025211663,0.0014479938],"genre_scores_gemma":[0.9418766,0.00018747139,0.053984042,0.00004007218,0.000019189107,0.00011435104,0.0031176417,0.000117923446,0.0005427961],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99919623,0.00024276401,0.000048337643,0.00030555663,0.00014037339,0.00006673623],"domain_scores_gemma":[0.99248064,0.0052343607,0.00069184194,0.0005893546,0.00070129713,0.00030250865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001113899,0.0009767207,0.00040118856,0.0028646511,0.00035211002,0.00075202907,0.0006653769,0.00088498107,0.0006475676],"category_scores_gemma":[0.013063849,0.00039737992,0.0007844701,0.0012980872,0.00023039136,0.0019331862,0.00088448497,0.0011577443,0.0006528044],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005184054,0.0005760046,0.26319298,0.00038408075,0.0002562116,0.00028950936,0.00074325956,0.39327848,0.00863561,0.0020670884,0.0066637173,0.3233947],"study_design_scores_gemma":[0.000011619443,0.000058986683,0.015262033,0.000017660852,0.00002437372,0.000044920216,0.0000739949,0.97959125,0.0016400613,0.0025628244,0.0006990604,0.00001325027],"about_ca_topic_score_codex":0.007698805,"about_ca_topic_score_gemma":0.020136021,"teacher_disagreement_score":0.007698805,"about_ca_system_score_codex":0.0007374711,"about_ca_system_score_gemma":0.0008305827,"threshold_uncertainty_score":0.015307963},"labels":[],"label_agreement":null},{"id":"W3122444274","doi":"10.7287/peerj.preprints.1920v2","title":"Analysis of Test Driven Development on sentiment and coding activities in GitHub repositories","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Java; Computer science; Coding (social sciences); Test-driven development; Set (abstract data type); Software; Control (management); Software engineering; Software development; World Wide Web; Operating system; Programming language; Artificial intelligence","score_opus":0.01393982838113878,"score_gpt":0.2520172469555547,"score_spread":0.23807741857441594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122444274","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9979366,0.00006755531,0.00077151216,0.000065487344,0.0000054183615,0.000021109507,0.0003518419,0.000037941893,0.0007424087],"genre_scores_gemma":[0.99753296,0.00006140495,0.0010462556,0.00002477974,0.000013579731,0.000030363117,0.0009406491,0.000020463503,0.000329572],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.995141,0.0015623778,0.00044110912,0.0003477778,0.0021276043,0.00038011788],"domain_scores_gemma":[0.9426059,0.029916571,0.014031223,0.0017336966,0.010563905,0.0011487714],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0038954497,0.00029058108,0.00041702206,0.0045439736,0.0003417825,0.0012507823,0.00036686365,0.0002962013,0.0005142069],"category_scores_gemma":[0.03528881,0.00016523332,0.00034161945,0.0040374044,0.00040011166,0.0010092335,0.0011076444,0.00040011606,0.00018635526],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053929136,0.00022202065,0.8990643,0.0002903823,0.00012220684,0.0004538293,0.0049845204,0.00094826886,0.010623547,0.00030613586,0.0015321318,0.08091345],"study_design_scores_gemma":[0.0000074334143,0.0001754004,0.98653525,0.000042990316,0.000036634392,0.00027600434,0.0027545767,0.005568785,0.003020413,0.00014330447,0.0014109997,0.00002813293],"about_ca_topic_score_codex":0.002857902,"about_ca_topic_score_gemma":0.0027216699,"teacher_disagreement_score":0.99610454,"about_ca_system_score_codex":0.0008741361,"about_ca_system_score_gemma":0.00047826464,"threshold_uncertainty_score":0.020601332},"labels":[],"label_agreement":null},{"id":"W3122526936","doi":"10.1109/bigdse.2015.14","title":"Software Analytics to Software Practice: A Systematic Literature Review","year":2015,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Software analytics; Artifact (error); Computer science; Data science; Software; Analytics; Systematic review; Process (computing); Big data; Software engineering; Field (mathematics); Software development; Software development process; Data mining; Artificial intelligence","score_opus":0.03346110073823944,"score_gpt":0.3473474331374403,"score_spread":0.3138863323992009,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122526936","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027914552,0.9906097,0.0014831543,0.001536697,0.0003108303,0.0015308025,0.0006979097,0.000025771722,0.0010136317],"genre_scores_gemma":[0.024200788,0.9649876,0.005685111,0.0011749496,0.00010821868,0.0031229325,0.0005167642,0.000019676085,0.00018389456],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9678876,0.013261505,0.010494693,0.0018608692,0.005818481,0.0006768581],"domain_scores_gemma":[0.8774881,0.08958658,0.011070727,0.0023638348,0.018210735,0.0012800504],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026729656,0.0018427611,0.0058700256,0.043866023,0.0018799378,0.004444019,0.002359279,0.0028684144,0.004006332],"category_scores_gemma":[0.12946734,0.0018193576,0.0046234927,0.033856113,0.0019524966,0.006379365,0.0040515875,0.0020302315,0.0006273471],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009482709,0.000051533883,0.001626338,0.86949396,0.0020859276,0.0004251613,0.0026779869,0.00026425626,0.00036443234,0.0017340599,0.0052331053,0.11594851],"study_design_scores_gemma":[0.000053233045,0.00010146948,0.0023608543,0.9555714,0.0057833986,0.00035793858,0.0025741954,0.00011695315,0.00016568774,0.0010044813,0.031872194,0.000038211092],"about_ca_topic_score_codex":0.010775527,"about_ca_topic_score_gemma":0.03132247,"teacher_disagreement_score":0.97327036,"about_ca_system_score_codex":0.0096946545,"about_ca_system_score_gemma":0.042065818,"threshold_uncertainty_score":0.14136165},"labels":[],"label_agreement":null},{"id":"W3122715434","doi":"10.7287/peerj.preprints.2373","title":"Stopping duplicate bug reports before they start with Continuous Querying for bug reports","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Data deduplication; Search engine indexing; BitTorrent tracker; Software bug; Security bug; Information retrieval; Process (computing); Software; Database; Artificial intelligence; Programming language; Cloud computing","score_opus":0.01488962674465569,"score_gpt":0.24715986460304953,"score_spread":0.23227023785839385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122715434","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31673542,0.0051596374,0.6229573,0.002951275,0.0006053594,0.0017463678,0.0015559989,0.040889516,0.007399052],"genre_scores_gemma":[0.5562947,0.00083929865,0.43310174,0.0007578951,0.00030400485,0.00031397014,0.0021443078,0.0019985393,0.004245578],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97844726,0.005896398,0.003069079,0.0033494625,0.008471,0.00076684507],"domain_scores_gemma":[0.80676454,0.10069859,0.026022708,0.042434957,0.021548502,0.0025306565],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021234063,0.001658433,0.0029882593,0.0059087775,0.0012342358,0.0047993283,0.003799966,0.0022436674,0.0025887326],"category_scores_gemma":[0.12853365,0.0009961155,0.0011025948,0.003675981,0.0016097212,0.0076442715,0.0036611976,0.002557806,0.0019871641],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014544276,0.00091457623,0.05702697,0.0019069434,0.00032405442,0.00064708886,0.0043186555,0.0054660067,0.051351078,0.00611398,0.013862133,0.8566141],"study_design_scores_gemma":[0.0011522331,0.008047968,0.13977969,0.0014460685,0.0016743157,0.009898221,0.007600933,0.3201671,0.35091066,0.03978027,0.11821634,0.0013261674],"about_ca_topic_score_codex":0.0018352949,"about_ca_topic_score_gemma":0.0018246147,"teacher_disagreement_score":0.021234063,"about_ca_system_score_codex":0.00093931094,"about_ca_system_score_gemma":0.0026234565,"threshold_uncertainty_score":0.11229783},"labels":[],"label_agreement":null},{"id":"W3122939027","doi":"10.7287/peerj.preprints.1705","title":"The unreasonable effectiveness of traditional information retrieval in crash report deduplication","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Crash; Computer science; Data deduplication; Software; Scalability; Information retrieval; Precision and recall; Set (abstract data type); Database; Data science; Data mining; Software engineering; World Wide Web; Operating system","score_opus":0.015265970731580713,"score_gpt":0.2492900997646688,"score_spread":0.23402412903308809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122939027","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53172815,0.04095521,0.35745066,0.006870482,0.0020930173,0.0015390507,0.00619617,0.0314094,0.02175775],"genre_scores_gemma":[0.6811616,0.006312379,0.29667437,0.0013489512,0.00058661326,0.00029327418,0.005778881,0.0009282769,0.0069156364],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9896402,0.002831167,0.0012938315,0.0017520989,0.003981525,0.0005011604],"domain_scores_gemma":[0.9397359,0.033910356,0.0027656755,0.016785968,0.0062937094,0.00050835265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01226868,0.0014046076,0.0019198285,0.0059565324,0.0018534528,0.0045128483,0.0031991636,0.0021277403,0.0017409248],"category_scores_gemma":[0.055410847,0.00069924956,0.00092993944,0.0068978905,0.0017088328,0.010702564,0.002119615,0.001726293,0.0037017548],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013888711,0.00058770826,0.014645534,0.002703605,0.00044814264,0.0004463361,0.0010649869,0.022437092,0.040026277,0.0048330175,0.05123869,0.86017966],"study_design_scores_gemma":[0.0007431934,0.003742386,0.046286874,0.0010599336,0.001165324,0.008314438,0.0054498855,0.37656668,0.37691513,0.037632488,0.14132991,0.0007937249],"about_ca_topic_score_codex":0.004745031,"about_ca_topic_score_gemma":0.0062389793,"teacher_disagreement_score":0.01226868,"about_ca_system_score_codex":0.0013932934,"about_ca_system_score_gemma":0.0022068669,"threshold_uncertainty_score":0.06488377},"labels":[],"label_agreement":null},{"id":"W3123088576","doi":"10.1007/s10664-020-09900-0","title":"Investigating design anti-pattern and design pattern mutations and their change- and fault-proneness","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Computer Research Institute of Montréal; Polytechnique Montréal","funders":"","keywords":"Software design pattern; Structural pattern; Software design; Computer science; Design pattern; Software evolution; Software quality; Software; Software development; Software engineering; Software construction; Programming language","score_opus":0.09046305438679865,"score_gpt":0.2872581248634783,"score_spread":0.19679507047667966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3123088576","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9973074,0.000071233495,0.0020202235,0.000050888575,0.0000022507375,0.00001320583,0.00006699455,0.000021993048,0.0004457995],"genre_scores_gemma":[0.99873966,0.000017324997,0.0010091352,0.000007415255,0.0000013618836,0.0000087193175,0.000059410053,0.0000069467155,0.00015016996],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9945041,0.002355385,0.0004960385,0.0009199658,0.0014572084,0.00026723056],"domain_scores_gemma":[0.75764173,0.1747767,0.04703725,0.012071723,0.0071450467,0.0013275591],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067707193,0.0003487314,0.00026780204,0.0018849864,0.0002778272,0.0009966785,0.000857642,0.0007969543,0.0019624934],"category_scores_gemma":[0.11432311,0.000292898,0.00041885776,0.0014342013,0.0008645757,0.0017827207,0.0007298924,0.0011461513,0.00019506425],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044379753,0.0004348337,0.9511425,0.00007754081,0.00023275215,0.00021592999,0.001185961,0.0051816925,0.00511491,0.0017062429,0.00013457665,0.03412922],"study_design_scores_gemma":[0.00005978881,0.00073955965,0.93201226,0.00003866604,0.0001976011,0.0008660775,0.001786676,0.051741946,0.0064802305,0.0053383266,0.000703149,0.000035716414],"about_ca_topic_score_codex":0.0011074498,"about_ca_topic_score_gemma":0.0019765839,"teacher_disagreement_score":0.0067707193,"about_ca_system_score_codex":0.00060659985,"about_ca_system_score_gemma":0.00066280446,"threshold_uncertainty_score":0.03580737},"labels":[],"label_agreement":null},{"id":"W3124091587","doi":"10.7287/peerj.preprints.1260v1","title":"Comments on \"Researcher bias: The use of machine learning in software defect prediction\"","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; McGill University","funders":"","keywords":"Metric (unit); Computer science; Group (periodic table); Construct (python library); Predictive modelling; Reuse; Association (psychology); Machine learning; Selection (genetic algorithm); Artificial intelligence; Software; Model selection; Data mining; Data science; Econometrics; Psychology; Mathematics; Engineering","score_opus":0.1805958843775091,"score_gpt":0.31951372142697376,"score_spread":0.13891783704946467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3124091587","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005783431,0.00082390587,0.00073968316,0.97779137,0.019024532,0.000020197343,0.00014855615,0.00008792465,0.0007854387],"genre_scores_gemma":[0.005533479,0.0008754188,0.0009357376,0.9712315,0.01914343,0.000072813695,0.000060724855,0.00009943146,0.0020474724],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.96070015,0.016589642,0.0042732405,0.0041298335,0.012663698,0.0016434637],"domain_scores_gemma":[0.71041924,0.19737366,0.014084702,0.007882176,0.063781224,0.0064590024],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04418568,0.0021232853,0.0014458796,0.002245393,0.0064070597,0.0063240416,0.0069942214,0.035624158,0.006933502],"category_scores_gemma":[0.22961593,0.0012023854,0.0022862453,0.0029089416,0.0091473125,0.008983624,0.0050788755,0.04276887,0.005697603],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000468795,0.000018627665,0.0008445697,0.00013019165,0.000021369804,0.0003271676,0.0011502195,0.00016596673,0.00022271775,0.0026603877,0.9894482,0.0049636387],"study_design_scores_gemma":[0.000086111264,0.00011566384,0.0030812924,0.0013286115,0.00006955451,0.0011666637,0.006265502,0.0012059088,0.0013249668,0.009554461,0.97551006,0.000291294],"about_ca_topic_score_codex":0.015451428,"about_ca_topic_score_gemma":0.012865619,"teacher_disagreement_score":0.9558143,"about_ca_system_score_codex":0.0058942335,"about_ca_system_score_gemma":0.008555948,"threshold_uncertainty_score":0.233679},"labels":[],"label_agreement":null},{"id":"W3124432101","doi":"","title":"Applying empirical software engineering to software architecture: challenges and lessons learned","year":2010,"lang":"en","type":"article","venue":"Cineca Institutional Research Information System (Tor Vergata University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Software engineering; Resource-oriented architecture; Software peer review; Reference architecture; Computer science; Software development; Social software engineering; Software architecture; Software architecture description; Software construction; Architecture tradeoff analysis method; Engineering; Software; Systems engineering","score_opus":0.08950162488329166,"score_gpt":0.31077143457798273,"score_spread":0.22126980969469107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3124432101","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05438698,0.15037186,0.17860131,0.59845334,0.0038137424,0.00053476513,0.00015585069,0.00026889038,0.013413323],"genre_scores_gemma":[0.5561656,0.17059503,0.22474426,0.039208278,0.0053421017,0.0012952433,0.00021367091,0.00020870173,0.0022271094],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.8258892,0.14417678,0.0071158144,0.0052428665,0.015811527,0.0017637382],"domain_scores_gemma":[0.34267884,0.59309596,0.007346323,0.02535497,0.028689215,0.0028346686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.17771734,0.0015940511,0.0024944511,0.004869732,0.0029722434,0.013987913,0.008177353,0.0062152725,0.0021006675],"category_scores_gemma":[0.34905535,0.0012609187,0.0012566474,0.0070007686,0.026047466,0.03172628,0.008755836,0.012048772,0.00058606145],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013210914,0.000941718,0.014112355,0.0075832717,0.00032040753,0.00085361215,0.020022558,0.0077470965,0.0007727011,0.35618645,0.01762953,0.57369816],"study_design_scores_gemma":[0.0001332518,0.0003499633,0.0049874573,0.009590124,0.00006594618,0.00044623952,0.04291798,0.010907766,0.0008962465,0.8689869,0.060580924,0.00013719875],"about_ca_topic_score_codex":0.0059703044,"about_ca_topic_score_gemma":0.006792685,"teacher_disagreement_score":0.17771734,"about_ca_system_score_codex":0.00716422,"about_ca_system_score_gemma":0.016277485,"threshold_uncertainty_score":0.93987036},"labels":[],"label_agreement":null},{"id":"W3124537160","doi":"10.7287/peerj.preprints.826","title":"An empirical study of goto in C code","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Goto; Go/no go; Commit; Computer science; Statement (logic); Code (set theory); Dijkstra's algorithm; Programming language; Limit (mathematics); Empirical research; Statistics; Mathematics; Theoretical computer science; Machine learning; Law; Database; Political science; Shortest path problem; Graph; Set (abstract data type)","score_opus":0.09214601773672157,"score_gpt":0.39471980448581767,"score_spread":0.3025737867490961,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3124537160","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9956371,0.00018075533,0.0007133651,0.00035318814,0.00000916307,0.000059162212,0.00013235358,0.000026697162,0.0028883382],"genre_scores_gemma":[0.99768853,0.00014888938,0.00090737175,0.00014513786,0.000011612674,0.000050917395,0.00023066836,0.000045587196,0.0007711233],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98542494,0.005386125,0.0010256529,0.00169572,0.0055935206,0.0008740806],"domain_scores_gemma":[0.6235436,0.2614383,0.059109587,0.014926677,0.03541363,0.0055682887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011253402,0.00032390738,0.00031493598,0.003170518,0.0020316106,0.0024762505,0.0012658125,0.0011250505,0.0021706696],"category_scores_gemma":[0.17576379,0.0005431214,0.00026887655,0.003939009,0.003921851,0.0049894755,0.0019500674,0.0025119937,0.0005158586],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003929659,0.0005919943,0.90455854,0.00046189944,0.00006938513,0.00078958704,0.046479523,0.00044340696,0.0017100341,0.001977114,0.0020542115,0.040471457],"study_design_scores_gemma":[0.000040083716,0.00082339527,0.92378134,0.00046639398,0.000051848296,0.0011329413,0.05551569,0.00336977,0.002080125,0.0013379971,0.011305911,0.00009443397],"about_ca_topic_score_codex":0.00845245,"about_ca_topic_score_gemma":0.011551135,"teacher_disagreement_score":0.011253402,"about_ca_system_score_codex":0.0016390551,"about_ca_system_score_gemma":0.0017038301,"threshold_uncertainty_score":0.059514403},"labels":[],"label_agreement":null},{"id":"W3124682540","doi":"10.7287/peerj.preprints.1771","title":"Judging a commit by its cover; or can a commit message predict build failure?","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Commit; Computer science; Code (set theory); Source code; Proxy (statistics); Database; Programming language; Computer security; Machine learning","score_opus":0.013777227243916359,"score_gpt":0.245285266788623,"score_spread":0.23150803954470664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3124682540","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97257894,0.00026731216,0.021351412,0.0009894891,0.000075060874,0.00004311707,0.0010754891,0.00062665873,0.002992531],"genre_scores_gemma":[0.99500704,0.000038608225,0.0036991348,0.000069409965,0.000039786704,0.000014014639,0.00070172275,0.000064641805,0.00036570217],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99590826,0.0010611244,0.00046739882,0.00078182877,0.0013549083,0.00042653747],"domain_scores_gemma":[0.90584624,0.057137102,0.01740227,0.006847434,0.009736785,0.003030185],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0066422173,0.0006027412,0.0006151563,0.0030763904,0.00054208783,0.0018643804,0.00064974866,0.0011392946,0.0016364409],"category_scores_gemma":[0.087616146,0.0003250594,0.0003220313,0.0021180373,0.0010172059,0.0036862623,0.0015317388,0.0014603455,0.0011222867],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045290845,0.00010619148,0.914763,0.00016470441,0.00016740362,0.000115454044,0.0012250227,0.004501176,0.0052462723,0.00096102303,0.0033950582,0.068901755],"study_design_scores_gemma":[0.00002677884,0.00028969703,0.8790768,0.00010150515,0.00008344834,0.00028173218,0.0021197295,0.10169473,0.006087122,0.006521377,0.0035900103,0.00012697805],"about_ca_topic_score_codex":0.0038495534,"about_ca_topic_score_gemma":0.009296868,"teacher_disagreement_score":0.9933578,"about_ca_system_score_codex":0.00049681956,"about_ca_system_score_gemma":0.00069835945,"threshold_uncertainty_score":0.03512788},"labels":[],"label_agreement":null},{"id":"W3125018233","doi":"10.2308/iace-50144","title":"The Effect of Business Process Representation Type on Assessment of Business and Control Risks: Diagrams versus Narratives","year":2012,"lang":"en","type":"article","venue":"Issues in Accounting Education","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Diagrammatic reasoning; Representation (politics); Narrative; Control (management); Affect (linguistics); Process (computing); Task (project management); Computer science; Perception; Psychology; Accounting; Business; Artificial intelligence; Political science; Management; Economics; Linguistics","score_opus":0.024929231058918356,"score_gpt":0.40416119496461966,"score_spread":0.3792319639057013,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3125018233","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9958507,0.00007776184,0.0012390859,0.000120539524,0.00001688427,0.00006793157,0.000029443294,0.000033263354,0.0025643783],"genre_scores_gemma":[0.99714524,0.000069235015,0.0017813359,0.0000385872,0.000009610704,0.00008961395,0.000042052434,0.000017334227,0.00080702227],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9858483,0.010146805,0.0013004015,0.0007467959,0.0016726797,0.00028499568],"domain_scores_gemma":[0.54306924,0.4113318,0.03063518,0.0059899306,0.005597894,0.003376049],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011030787,0.00046639369,0.00045046845,0.0008508077,0.0003792974,0.004081577,0.00072809204,0.0008905224,0.007319088],"category_scores_gemma":[0.19851731,0.00033014454,0.00053060864,0.0006425483,0.0010018336,0.003220555,0.0015363247,0.0012918039,0.00051633874],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.044212136,0.022401664,0.52458066,0.0021578374,0.001067853,0.0006109707,0.033323202,0.01718871,0.058266643,0.010221507,0.0018137961,0.28415498],"study_design_scores_gemma":[0.0020585207,0.02822342,0.8288288,0.0014698963,0.0016386363,0.0006251356,0.022168403,0.05212218,0.042649563,0.011105603,0.008649382,0.00046039693],"about_ca_topic_score_codex":0.00073192874,"about_ca_topic_score_gemma":0.00056297146,"teacher_disagreement_score":0.011030787,"about_ca_system_score_codex":0.0006367781,"about_ca_system_score_gemma":0.00045872224,"threshold_uncertainty_score":0.058337092},"labels":[],"label_agreement":null},{"id":"W3125821053","doi":"10.7287/peerj.preprints.1138","title":"The charming code that error messages are talking about","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Python (programming language); Cyclomatic complexity; Computer science; Debugging; Programming language; Random testing; Code coverage; Software quality; Source lines of code; Software bug; Software; Syntax error; Software metric; Charm (quantum number); Code (set theory); Source code; Algorithm; Abstract syntax tree; Test case; Software development; Particle physics; Machine learning; Physics","score_opus":0.07207317227592544,"score_gpt":0.3135958871748,"score_spread":0.24152271489887459,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3125821053","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29709652,0.0061115306,0.41634974,0.027516391,0.009213745,0.0012314975,0.006745044,0.0630041,0.17273147],"genre_scores_gemma":[0.667053,0.002284665,0.18001258,0.012492163,0.0015998723,0.0007260799,0.0037412816,0.016190862,0.11589956],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99168605,0.002464841,0.00059205794,0.0008665065,0.0039257654,0.00046480598],"domain_scores_gemma":[0.950754,0.019820156,0.010392356,0.007806764,0.010174995,0.001051796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026141615,0.0015769986,0.00058862503,0.0024870678,0.0020595687,0.0028588956,0.000929551,0.0022115104,0.013503694],"category_scores_gemma":[0.046489693,0.00055378984,0.00048277265,0.002005445,0.0034545278,0.0053307414,0.0027597174,0.003051101,0.006618674],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015626523,0.0003344461,0.058407843,0.0029268577,0.00027436137,0.0040674014,0.01627849,0.0043713422,0.047951967,0.1461075,0.21124153,0.5064757],"study_design_scores_gemma":[0.00007202816,0.00040193554,0.02924158,0.0022151973,0.00022702881,0.0064423615,0.0035824536,0.010631853,0.06547484,0.05838846,0.8229501,0.00037212577],"about_ca_topic_score_codex":0.0021918914,"about_ca_topic_score_gemma":0.0018684452,"teacher_disagreement_score":0.013503694,"about_ca_system_score_codex":0.0012060588,"about_ca_system_score_gemma":0.0016376926,"threshold_uncertainty_score":0.04517436},"labels":[],"label_agreement":null},{"id":"W3127656560","doi":"10.1109/tse.2021.3055123","title":"Deprecation of Packages and Releases in Software Ecosystems: A Case Study on NPM","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Notation; Software; Code (set theory); Programming language; Software engineering; Information retrieval; World Wide Web; Arithmetic; Mathematics; Set (abstract data type)","score_opus":0.01918480495773581,"score_gpt":0.2594253268007716,"score_spread":0.24024052184303582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3127656560","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9888692,0.00033993102,0.0044751046,0.0010546488,0.000019541672,0.00006891584,0.00009180558,0.00014334216,0.004937492],"genre_scores_gemma":[0.9869372,0.00033417562,0.009584881,0.00030396206,0.000033710232,0.000084987645,0.00019971652,0.00010708912,0.0024141206],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9937516,0.0031827455,0.00041505918,0.0006569414,0.001409376,0.000584203],"domain_scores_gemma":[0.95123416,0.03320478,0.00726484,0.0039366516,0.0025221591,0.0018373859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007026562,0.00036306636,0.00030204875,0.00227705,0.003939587,0.0023092753,0.0013859707,0.002055843,0.0015213796],"category_scores_gemma":[0.029985579,0.00040909316,0.0005817636,0.002847458,0.002733676,0.005051306,0.0029667756,0.0018716092,0.00044178448],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005279112,0.0019043346,0.5112999,0.0008824654,0.00010859475,0.056217916,0.21023716,0.0061808433,0.010097101,0.025628503,0.0108716665,0.16604364],"study_design_scores_gemma":[0.00013327113,0.0010506245,0.5102687,0.0009857701,0.00017439065,0.028574906,0.2378517,0.041041423,0.012131937,0.014304724,0.15317412,0.00030839708],"about_ca_topic_score_codex":0.009061076,"about_ca_topic_score_gemma":0.0141529655,"teacher_disagreement_score":0.009061076,"about_ca_system_score_codex":0.0024953648,"about_ca_system_score_gemma":0.0021182187,"threshold_uncertainty_score":0.037160516},"labels":[],"label_agreement":null},{"id":"W3128185208","doi":"10.1145/3437479.3437484","title":"Summary of the 2nd International Workshop on Bots in Software Engineering (BotSE 2020)","year":2021,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Software engineering; Presentation (obstetrics); Computer science; Software; World Wide Web; Engineering management; Engineering; Programming language","score_opus":0.015559303008806358,"score_gpt":0.2501550741661026,"score_spread":0.23459577115729624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3128185208","genre_codex":"editorial","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013806141,0.054441188,0.08543375,0.099910036,0.407918,0.0032990552,0.017003212,0.008497387,0.30969128],"genre_scores_gemma":[0.030174425,0.030856619,0.030260742,0.01997346,0.04700873,0.0021373217,0.038624413,0.0048220675,0.7961422],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99622667,0.0007164476,0.00020734253,0.0006846733,0.0015186317,0.00064623327],"domain_scores_gemma":[0.98932797,0.0009782033,0.0002465951,0.0005320933,0.0044773794,0.0044376943],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0071856673,0.0022987078,0.0011877635,0.0029471866,0.0021020533,0.008634585,0.0021237633,0.003131288,0.08551286],"category_scores_gemma":[0.0073902984,0.0007511642,0.0016373611,0.0024430277,0.0006145759,0.0062988643,0.0068331617,0.0051180054,0.06360612],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013287172,0.00017618145,0.00023339469,0.0004190649,0.000016567981,0.00011291031,0.00022626376,0.0002929305,0.0013344687,0.0014754327,0.93509156,0.060488295],"study_design_scores_gemma":[0.00003313988,0.00014761269,0.0008101378,0.0003212035,0.000018105446,0.00009711719,0.00030110372,0.00041631603,0.0007334322,0.0013262478,0.9957599,0.000035700712],"about_ca_topic_score_codex":0.0024219195,"about_ca_topic_score_gemma":0.0056563457,"teacher_disagreement_score":0.99281436,"about_ca_system_score_codex":0.0017163492,"about_ca_system_score_gemma":0.0032940516,"threshold_uncertainty_score":0.28606904},"labels":[],"label_agreement":null},{"id":"W3128560401","doi":"10.1109/icse-seip52600.2021.00044","title":"Refactoring Practices in the Context of Modern Code Review: An Industrial Case Study at Xerox","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec","funders":"","keywords":"Code refactoring; Documentation; Codebase; Code review; Computer science; Software engineering; Context (archaeology); Best practice; Software; Software quality; Software development; Programming language","score_opus":0.2502142211954906,"score_gpt":0.4137484406022141,"score_spread":0.16353421940672352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3128560401","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99156976,0.0005397045,0.0044373823,0.0011399946,0.000019864023,0.00021232509,0.00004227212,0.00006586532,0.0019727228],"genre_scores_gemma":[0.9850038,0.0006243823,0.011585294,0.0003843609,0.00004384458,0.00018720448,0.000092377784,0.00007591705,0.0020027077],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9658693,0.023272574,0.0017287758,0.0022221273,0.005151199,0.0017560519],"domain_scores_gemma":[0.7862303,0.15304282,0.018719075,0.009530112,0.02557447,0.0069032516],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02861027,0.0004705379,0.0005257825,0.0031506154,0.0058015366,0.0032323494,0.0018398084,0.0025871645,0.0012553583],"category_scores_gemma":[0.07953603,0.0006923226,0.00040545085,0.0026257944,0.0029707823,0.0030710017,0.0029295636,0.0020829213,0.00045296978],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007466375,0.00437509,0.16447541,0.0015048212,0.00009448806,0.033736408,0.548207,0.0033575837,0.013299261,0.0047857948,0.008540638,0.21687691],"study_design_scores_gemma":[0.00047554076,0.006316466,0.2649307,0.0019431235,0.00022630609,0.027696224,0.52649397,0.018565694,0.02257318,0.0063265967,0.123814784,0.00063747243],"about_ca_topic_score_codex":0.00796969,"about_ca_topic_score_gemma":0.018584793,"teacher_disagreement_score":0.02861027,"about_ca_system_score_codex":0.004148335,"about_ca_system_score_gemma":0.0052367523,"threshold_uncertainty_score":0.1513074},"labels":[],"label_agreement":null},{"id":"W3128670333","doi":"10.5815/ijitcs.2021.01.01","title":"Duration Estimation Models for Open Source Software Projects","year":2021,"lang":"en","type":"article","venue":"International Journal of Information Technology and Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Duration (music); Estimation; Computer science; Swift; Software; Variable (mathematics); Government (linguistics); Open source software; Open source; Software engineering; Database; Operations research; Operating system; Systems engineering; Engineering","score_opus":0.016543310594426115,"score_gpt":0.28586069662440466,"score_spread":0.26931738602997857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3128670333","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22370759,0.003931667,0.761315,0.0014838175,0.00019114208,0.00035254628,0.0025965974,0.000697547,0.005724122],"genre_scores_gemma":[0.897662,0.0027308755,0.07229235,0.00021011519,0.00031546722,0.0010171011,0.006624438,0.00026169166,0.01888592],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965814,0.0015391423,0.00023810448,0.000814626,0.00043426835,0.0003925776],"domain_scores_gemma":[0.96756715,0.025215821,0.0034481473,0.00096255797,0.0020736598,0.0007326729],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011731431,0.0016253238,0.001661076,0.0027577365,0.00071055966,0.0023547078,0.0030000526,0.0020381787,0.00682358],"category_scores_gemma":[0.03252839,0.00093046186,0.0018872154,0.0023390218,0.00091577996,0.0026177932,0.0021650512,0.0033207973,0.0016039235],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047419933,0.00025537677,0.034593426,0.00033480942,0.00030580733,0.00036402204,0.0013310225,0.8317544,0.0008139785,0.05160557,0.0041128374,0.074054524],"study_design_scores_gemma":[0.000021786775,0.00005699298,0.0041185874,0.00005615832,0.000050897113,0.000046147547,0.00014586572,0.98036814,0.00013237953,0.013108893,0.0018666824,0.00002750589],"about_ca_topic_score_codex":0.020191202,"about_ca_topic_score_gemma":0.010178326,"teacher_disagreement_score":0.020191202,"about_ca_system_score_codex":0.002137835,"about_ca_system_score_gemma":0.0016786987,"threshold_uncertainty_score":0.062042475},"labels":[],"label_agreement":null},{"id":"W3129150043","doi":"10.1007/s10664-020-09902-y","title":"On the Removal of Feature Toggles","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Queen's University","funders":"","keywords":"Python (programming language); Feature (linguistics); Computer science; Source code; Cyclomatic complexity; Software; Data science; Software engineering; Programming language","score_opus":0.024005836158592978,"score_gpt":0.27051804039179844,"score_spread":0.24651220423320547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3129150043","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26561174,0.0020523726,0.65127873,0.0123526985,0.0014758633,0.0002657575,0.000935379,0.006134538,0.05989295],"genre_scores_gemma":[0.7237977,0.00082405447,0.24125686,0.0018552797,0.00054013444,0.00010161232,0.0015153444,0.001954973,0.02815414],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99486214,0.0016556684,0.00020357058,0.0006421348,0.0019867592,0.0006498408],"domain_scores_gemma":[0.94535077,0.026728928,0.0024236925,0.018890647,0.005873825,0.00073203276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005107956,0.0007893287,0.0011921955,0.0022007392,0.0018562141,0.0026756094,0.0027292583,0.0024565568,0.009828616],"category_scores_gemma":[0.059493095,0.0005290792,0.0010958998,0.0020065983,0.0021852176,0.005316001,0.0030775152,0.0023810987,0.003163823],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010431948,0.00059788435,0.013332577,0.00034198124,0.00014879623,0.00070127245,0.0006635691,0.017035166,0.013023469,0.07428856,0.030526733,0.84829676],"study_design_scores_gemma":[0.0004386247,0.000790369,0.052335396,0.00053235324,0.0006387322,0.0023721054,0.0022817615,0.41372076,0.04836959,0.37668535,0.10156377,0.0002711979],"about_ca_topic_score_codex":0.0059734234,"about_ca_topic_score_gemma":0.009671143,"teacher_disagreement_score":0.009828616,"about_ca_system_score_codex":0.0006545856,"about_ca_system_score_gemma":0.002315817,"threshold_uncertainty_score":0.03288001},"labels":[],"label_agreement":null},{"id":"W3130266246","doi":"","title":"An exploratory study on the introduction and removal of different types of technical debt in deep learning frameworks","year":2021,"lang":"en","type":"article","venue":"Institutional Knowledge (InK) - Institutional Knowledge at Singapore Management University (Singapore Management University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Technical debt; Debt; Computer science; Software; Business; Software development; Finance","score_opus":0.015261273800941243,"score_gpt":0.2309297704452297,"score_spread":0.21566849664428844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3130266246","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99363106,0.000075398064,0.0008207965,0.0002297075,0.000007707181,0.000052104133,0.000029169962,0.000014000506,0.005139976],"genre_scores_gemma":[0.99758255,0.00004246612,0.00096889253,0.00009018209,0.0000048531397,0.000029189543,0.000051395233,0.000011548922,0.0012187896],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99495655,0.0024061028,0.0002583875,0.00040133315,0.0011636424,0.0008140412],"domain_scores_gemma":[0.93495655,0.046545766,0.008286201,0.0034577558,0.003627829,0.0031258902],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008222788,0.00019363071,0.0002615701,0.0008467303,0.0017244614,0.002727936,0.0014591049,0.0014080844,0.003041316],"category_scores_gemma":[0.05231935,0.0002614741,0.00023298452,0.0010828767,0.0015135492,0.004671659,0.002256435,0.002990785,0.00025454685],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021749781,0.00980646,0.4036848,0.0010158592,0.00014922873,0.0026708893,0.22490719,0.0036918358,0.017887032,0.048361253,0.004291235,0.28135931],"study_design_scores_gemma":[0.00024372993,0.003979773,0.54462975,0.00095912156,0.00022419209,0.0012690106,0.31823197,0.02180082,0.013556271,0.018646386,0.076258875,0.00020010656],"about_ca_topic_score_codex":0.0039386535,"about_ca_topic_score_gemma":0.0077024777,"teacher_disagreement_score":0.008222788,"about_ca_system_score_codex":0.0025168005,"about_ca_system_score_gemma":0.00257605,"threshold_uncertainty_score":0.043486774},"labels":[],"label_agreement":null},{"id":"W3131995106","doi":"10.1145/3434279","title":"Are Comments on Stack Overflow Well Organized for Easy Retrieval by Developers?","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Manitoba; Concordia University; Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Classifier (UML); Information retrieval; Obsolescence; Artificial intelligence; Machine learning; Mechanism (biology); Point (geometry); Data mining; Data science","score_opus":0.08626475335032763,"score_gpt":0.32677058681156684,"score_spread":0.24050583346123922,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3131995106","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95414984,0.0010583075,0.027941374,0.0018144443,0.00016429379,0.0005900162,0.0015493885,0.0053802677,0.007352117],"genre_scores_gemma":[0.96641535,0.00037965467,0.028202614,0.00043081358,0.00014568446,0.00016635658,0.0015784225,0.00043835747,0.002242705],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99268544,0.0027330131,0.00082874787,0.0007781417,0.0022956517,0.0006789934],"domain_scores_gemma":[0.90665025,0.05079152,0.01810075,0.005378583,0.016783798,0.0022950885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007877054,0.00087728014,0.00086481794,0.0039783437,0.0011950207,0.0024300192,0.00073401234,0.0015127467,0.0026725864],"category_scores_gemma":[0.09273085,0.00039593625,0.0005213152,0.0017518863,0.0008792212,0.0056396085,0.0012546009,0.00078525464,0.0022183203],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024403136,0.00057979196,0.38659722,0.002963152,0.0002156288,0.0014553189,0.013406023,0.0031702227,0.06464763,0.0022423856,0.02927184,0.49301055],"study_design_scores_gemma":[0.0005214686,0.0032540804,0.6124345,0.0021709362,0.0009306206,0.004955551,0.027960539,0.16170827,0.09233421,0.0082200905,0.08477211,0.0007376615],"about_ca_topic_score_codex":0.004639848,"about_ca_topic_score_gemma":0.0050940607,"teacher_disagreement_score":0.007877054,"about_ca_system_score_codex":0.0008575472,"about_ca_system_score_gemma":0.0020455886,"threshold_uncertainty_score":0.041658342},"labels":[],"label_agreement":null},{"id":"W3132324754","doi":"10.3758/s13415-021-00872-2","title":"Correction to: Recallable but not Recognizable: The Influence of Semantic Priming in Recall Paradigms","year":2021,"lang":"en","type":"article","venue":"Cognitive Affective & Behavioral Neuroscience","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Wilfrid Laurier University","funders":"","keywords":"Recall; Priming (agriculture); Natural language processing; Computer science; Artificial intelligence; Linguistics; Psychology; Information retrieval; Cognitive psychology; Philosophy; Biology","score_opus":0.041291253512697156,"score_gpt":0.3188554248775399,"score_spread":0.27756417136484274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3132324754","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00053116336,0.00057489006,0.00088654814,0.06618588,0.92702,0.000047497415,0.0014195676,0.00070141995,0.0026330969],"genre_scores_gemma":[0.06604232,0.0052728765,0.0086459825,0.1256804,0.5401146,0.00057195773,0.0030178265,0.003925993,0.24672809],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946061,0.0007023336,0.0013725574,0.00095096166,0.0017647339,0.0006033531],"domain_scores_gemma":[0.9173493,0.025135009,0.004532336,0.009102471,0.040601414,0.003279508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040172394,0.0024791982,0.0035138484,0.004330295,0.0044738348,0.004720378,0.004778546,0.011885271,0.0854935],"category_scores_gemma":[0.11815392,0.0013779438,0.0015167703,0.0031971023,0.0042209476,0.0029915585,0.0024490047,0.013526179,0.046904024],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019731623,0.00001823279,0.00016857324,0.0002983438,0.00004267967,0.0007365617,0.00012273583,0.00006149237,0.0002198316,0.0015208191,0.988193,0.00842052],"study_design_scores_gemma":[0.00018563609,0.00005226438,0.003325632,0.00052451226,0.00008929379,0.0018432694,0.00033266124,0.0007789795,0.0017672762,0.005183726,0.9857744,0.00014235995],"about_ca_topic_score_codex":0.011781484,"about_ca_topic_score_gemma":0.010611066,"teacher_disagreement_score":0.0854935,"about_ca_system_score_codex":0.0045825196,"about_ca_system_score_gemma":0.0038819364,"threshold_uncertainty_score":0.28600425},"labels":[],"label_agreement":null},{"id":"W3132847876","doi":"10.1109/tse.2021.3060918","title":"Studying Duplicate Logging Statements and Their Relationships With Code Clones","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Logging; Code (set theory); Programming language; Database; Software engineering; Ecology","score_opus":0.040146525901182346,"score_gpt":0.26808989881282014,"score_spread":0.2279433729116378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3132847876","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.978735,0.001593537,0.015659062,0.00035713505,0.000039693783,0.00012756808,0.0005733741,0.0015626549,0.0013519215],"genre_scores_gemma":[0.9796464,0.00037735258,0.01726821,0.00015319706,0.00003636772,0.00008457329,0.001291452,0.00034240916,0.0008001925],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.985642,0.0022621385,0.0015284661,0.003847396,0.006058429,0.00066161255],"domain_scores_gemma":[0.71903384,0.16245653,0.073515505,0.017837621,0.02503091,0.002125563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068012765,0.0007919438,0.00056082755,0.006175816,0.0011506375,0.0014385126,0.0013360115,0.00119314,0.00085340213],"category_scores_gemma":[0.121963635,0.0006527671,0.00051418116,0.0043501677,0.0015931057,0.0033079986,0.0017347571,0.0014381203,0.00027223324],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001644948,0.00013768785,0.91564775,0.0005191847,0.00015832488,0.0025276488,0.0043940227,0.0019863038,0.005736636,0.0008434416,0.0018767336,0.06600791],"study_design_scores_gemma":[0.00006600395,0.0004777842,0.86952335,0.00056637963,0.00073232123,0.013273334,0.00652915,0.059672784,0.027502475,0.00473963,0.016708383,0.00020840047],"about_ca_topic_score_codex":0.0055264076,"about_ca_topic_score_gemma":0.009856518,"teacher_disagreement_score":0.0068012765,"about_ca_system_score_codex":0.0010048135,"about_ca_system_score_gemma":0.0015381376,"threshold_uncertainty_score":0.03596902},"labels":[],"label_agreement":null},{"id":"W3133232930","doi":"10.1049/iet-sen.2019.0384","title":"Investigating the information value of different sources of evidence of developers’ expertise for bug assignment in open‐source projects","year":2020,"lang":"en","type":"article","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Alberta Innovates - Technology Futures; Faculty of Graduate Studies and Research, University of Regina","keywords":"Open source; Value (mathematics); Computer science; Knowledge management; Open source software; Software engineering; Process management; Engineering; Programming language; Software; Machine learning","score_opus":0.11969056004223318,"score_gpt":0.31950892678434967,"score_spread":0.1998183667421165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133232930","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9789578,0.0026377125,0.012038307,0.0006058023,0.000022649441,0.00003762854,0.0023922226,0.00018816494,0.003119807],"genre_scores_gemma":[0.99160683,0.00036772556,0.005634231,0.00003489778,0.00003913793,0.000025351705,0.0020684828,0.000043294192,0.00017997663],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9855883,0.0064785676,0.001398496,0.0017269466,0.0042593693,0.0005483399],"domain_scores_gemma":[0.58999944,0.35736626,0.026556123,0.0106640905,0.012945048,0.0024689503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014782019,0.0005867072,0.0008505823,0.026366968,0.00065590855,0.0032154294,0.0009624092,0.0020523318,0.0010616817],"category_scores_gemma":[0.17959729,0.00033940095,0.000808817,0.015222616,0.0013970688,0.006834602,0.003019779,0.0015382349,0.00028588867],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010178853,0.00037588147,0.81271714,0.0015564193,0.0012978123,0.00048368316,0.003798711,0.0071351747,0.0047316295,0.0040388564,0.0027533546,0.16009355],"study_design_scores_gemma":[0.00011419075,0.0005864936,0.8897758,0.0005230798,0.0011878518,0.0013993852,0.0029078568,0.075241275,0.009783786,0.013029412,0.005249298,0.0002014809],"about_ca_topic_score_codex":0.002949435,"about_ca_topic_score_gemma":0.0036033,"teacher_disagreement_score":0.026366968,"about_ca_system_score_codex":0.0010974684,"about_ca_system_score_gemma":0.00073644257,"threshold_uncertainty_score":0.07817578},"labels":[],"label_agreement":null},{"id":"W3133304533","doi":"10.1109/tse.2021.3058985","title":"A Study of C/C++ Code Weaknesses on Stack Overflow","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal; University of Manitoba; Concordia University; Huawei Technologies (Canada)","funders":"","keywords":"Notation; Computer science; Code (set theory); Stack (abstract data type); Programming language; Mathematical notation; Software; Theoretical computer science; Mathematics; Arithmetic","score_opus":0.02446365343770417,"score_gpt":0.26541750701339983,"score_spread":0.24095385357569565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133304533","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9922402,0.0005698143,0.0029370307,0.00045204518,0.000017837297,0.000056199133,0.0007367576,0.0002998415,0.0026902852],"genre_scores_gemma":[0.9900119,0.000336239,0.005873804,0.00021560008,0.00003169704,0.00006389046,0.0017824121,0.00024317605,0.0014412568],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9954987,0.00092498004,0.00040661974,0.00096853345,0.0019182526,0.0002829363],"domain_scores_gemma":[0.890596,0.07029823,0.023795426,0.004198512,0.009617789,0.0014940938],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036640675,0.00064916775,0.00034769782,0.0051942198,0.0011014139,0.0015510353,0.00081411336,0.001386025,0.0019443635],"category_scores_gemma":[0.07751864,0.0003987584,0.00037578613,0.0036934654,0.0016212395,0.0047479793,0.0018592627,0.0013813666,0.00069284503],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043487293,0.0002294251,0.8528728,0.000893318,0.00013734985,0.0017544654,0.01693768,0.003099532,0.0076708524,0.0033475088,0.00724287,0.105379365],"study_design_scores_gemma":[0.000035016732,0.0004818621,0.87544817,0.0008333297,0.00019903976,0.0063511175,0.017257322,0.053402625,0.013318527,0.005326613,0.02713246,0.00021393882],"about_ca_topic_score_codex":0.005307372,"about_ca_topic_score_gemma":0.006695543,"teacher_disagreement_score":0.005307372,"about_ca_system_score_codex":0.0007584026,"about_ca_system_score_gemma":0.00087954383,"threshold_uncertainty_score":0.019377649},"labels":[],"label_agreement":null},{"id":"W3133366476","doi":"10.1145/3439769","title":"Automatic API Usage Scenario Documentation from Technical Q&A Sites","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Saskatchewan; University of Calgary","funders":"","keywords":"Documentation; Computer science; Internal documentation; World Wide Web; Code (set theory); Application programming interface; Java; Information retrieval; Database; Programming language; Software development; Software","score_opus":0.057065708543481634,"score_gpt":0.21589802329278307,"score_spread":0.15883231474930143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133366476","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82340425,0.0016834978,0.10235388,0.0010474006,0.00024499293,0.0010579138,0.031086028,0.028060546,0.011061478],"genre_scores_gemma":[0.7011745,0.0006709918,0.21824533,0.00015958342,0.00017042023,0.00095329294,0.07270943,0.0012181877,0.004698214],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953695,0.0014125817,0.00050448417,0.00090368895,0.0015857753,0.00022393119],"domain_scores_gemma":[0.9737897,0.012029456,0.0044498593,0.002201311,0.0067266845,0.00080306764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030392436,0.00094926497,0.00068862806,0.012538388,0.0009113763,0.0016950292,0.0007713796,0.0009573318,0.0015643053],"category_scores_gemma":[0.023393942,0.0004999104,0.0006297445,0.005528316,0.00027428698,0.0021780108,0.0016017823,0.0007979122,0.0020318967],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006442978,0.00062660815,0.1388471,0.0026771165,0.0002329872,0.0021350717,0.0066391756,0.005839636,0.033232156,0.0025319662,0.08788896,0.71870494],"study_design_scores_gemma":[0.00019466995,0.0007678301,0.35521138,0.0009745754,0.00028083718,0.0030397088,0.008506925,0.391549,0.05462568,0.0077452273,0.17670435,0.00039992758],"about_ca_topic_score_codex":0.0022405686,"about_ca_topic_score_gemma":0.00596169,"teacher_disagreement_score":0.012538388,"about_ca_system_score_codex":0.0006677,"about_ca_system_score_gemma":0.001138065,"threshold_uncertainty_score":0.016073227},"labels":[],"label_agreement":null},{"id":"W3133561796","doi":"","title":"EEF-CAS: An Effort Estimation Framework with Customizable Attribute Selection","year":2013,"lang":"en","type":"article","venue":"Scholarship@Western (Western University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Flexibility (engineering); Process (computing); Estimation; Personalization; Set (abstract data type); Data mining; Software; Selection (genetic algorithm); Industrial engineering; Machine learning; Systems engineering","score_opus":0.05334393255199717,"score_gpt":0.30427238633678294,"score_spread":0.2509284537847858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133561796","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0023329686,0.00007891604,0.99561656,0.00010622597,0.000013380186,0.00014954526,0.00014882478,0.0008568245,0.000696815],"genre_scores_gemma":[0.12411489,0.00016209735,0.873463,0.000103637874,0.00006091766,0.00046419003,0.0007272924,0.00011211111,0.0007918961],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99150825,0.00355248,0.0006381442,0.0014291434,0.0025143009,0.0003576786],"domain_scores_gemma":[0.9868565,0.0068800705,0.0013992102,0.0015685108,0.002947586,0.00034815332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013660286,0.0015482649,0.0012875649,0.0045260787,0.000832891,0.002601976,0.0035398116,0.0012262003,0.002727377],"category_scores_gemma":[0.025465928,0.0005788532,0.0019401376,0.0032586018,0.0011740526,0.003641261,0.0032307517,0.0017998394,0.0005853413],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036089716,0.0004878014,0.020932008,0.0004413707,0.00047648518,0.00038486466,0.00095621526,0.2554687,0.0027737806,0.10646646,0.009420394,0.6018311],"study_design_scores_gemma":[0.00005789808,0.00022014325,0.0037612699,0.00010830215,0.00010515471,0.00025381043,0.00018650596,0.9237912,0.001772277,0.058526326,0.011105068,0.0001120122],"about_ca_topic_score_codex":0.009214924,"about_ca_topic_score_gemma":0.00940559,"teacher_disagreement_score":0.013660286,"about_ca_system_score_codex":0.0014752862,"about_ca_system_score_gemma":0.0031086833,"threshold_uncertainty_score":0.07224333},"labels":[],"label_agreement":null},{"id":"W3133564273","doi":"10.1145/3479529","title":"\"@alex, this fixes #9\": Analysis of Referencing Patterns in Pull Request Discussions","year":2021,"lang":"en","type":"article","venue":"Proceedings of the ACM on Human-Computer Interaction","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Referent; Variety (cybernetics); Source code; Thread (computing); Common ground; World Wide Web; User interface; Software; Interface (matter); Information retrieval; Human–computer interaction; Programming language; Artificial intelligence; Psychology","score_opus":0.05804857565293148,"score_gpt":0.3381942710071979,"score_spread":0.2801456953542664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133564273","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9548907,0.0004465431,0.016228858,0.00081911066,0.00008844796,0.00035720842,0.01559155,0.0030955367,0.008481985],"genre_scores_gemma":[0.89650583,0.00041129146,0.042114407,0.0006309682,0.00008061048,0.00090290693,0.04279707,0.0012640207,0.0152929155],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99564135,0.0015276269,0.000451783,0.00079189986,0.0012329242,0.0003544268],"domain_scores_gemma":[0.95864546,0.028286323,0.004852649,0.0029153002,0.004238757,0.0010615285],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0040953234,0.00037139712,0.00028783237,0.0054582083,0.0014392636,0.0018576275,0.00093685655,0.0010387017,0.0034135054],"category_scores_gemma":[0.024903938,0.00026605924,0.00040906118,0.004130905,0.000793255,0.0030120153,0.0021611822,0.0007919684,0.0016904428],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001780741,0.00074972113,0.45278656,0.0029533117,0.00018782947,0.0029960715,0.13320857,0.0018861507,0.030510923,0.014092533,0.07538552,0.28346208],"study_design_scores_gemma":[0.00009223655,0.00043773185,0.6344867,0.00084161334,0.00013468484,0.001793212,0.08497528,0.021551928,0.02566457,0.0071265777,0.22262515,0.00027033256],"about_ca_topic_score_codex":0.0070600337,"about_ca_topic_score_gemma":0.012342684,"teacher_disagreement_score":0.9959047,"about_ca_system_score_codex":0.0009878027,"about_ca_system_score_gemma":0.0010484927,"threshold_uncertainty_score":0.02165842},"labels":[],"label_agreement":null},{"id":"W3133722238","doi":"10.1007/s11334-021-00388-5","title":"On the use of textual feature extraction techniques to support the automated detection of refactoring documentation","year":2021,"lang":"en","type":"article","venue":"Innovations in Systems and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Documentation; Maintainability; Commit; Artificial intelligence; Machine learning; Naive Bayes classifier; Software; Process (computing); Software documentation; Data mining; Software engineering; Programming language; Software development; Software development process; Database","score_opus":0.03321997824429876,"score_gpt":0.2960819489872761,"score_spread":0.26286197074297735,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133722238","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15905762,0.0014166192,0.81231576,0.0012715252,0.0001758961,0.0003428892,0.0012055481,0.016017474,0.008196711],"genre_scores_gemma":[0.32372218,0.0011039223,0.6651171,0.00044010562,0.00010122342,0.000117819545,0.0023480954,0.00089588587,0.006153679],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982901,0.0004164384,0.00016104047,0.00024483624,0.0007750242,0.00011244249],"domain_scores_gemma":[0.9822412,0.011777443,0.001044857,0.0018277715,0.002931854,0.00017690948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014991153,0.000656641,0.00054308336,0.0030968469,0.0006040674,0.002018717,0.001577848,0.0012721586,0.00269756],"category_scores_gemma":[0.01082798,0.00045337708,0.00074451056,0.0019068542,0.0005125453,0.0021900462,0.0010425464,0.00075813883,0.0016482531],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004369705,0.0003923648,0.0058486834,0.00027384088,0.00009237534,0.0005421041,0.00034913697,0.007982977,0.095756814,0.0036441828,0.0074903737,0.8771902],"study_design_scores_gemma":[0.00013192158,0.00046233658,0.013583841,0.00019582316,0.00021621515,0.0012469811,0.00034111657,0.74760604,0.20715833,0.010247248,0.018645164,0.00016493486],"about_ca_topic_score_codex":0.006224931,"about_ca_topic_score_gemma":0.009828506,"teacher_disagreement_score":0.006224931,"about_ca_system_score_codex":0.00035054854,"about_ca_system_score_gemma":0.00088444323,"threshold_uncertainty_score":0.012377381},"labels":[],"label_agreement":null},{"id":"W3134221463","doi":"10.1109/tse.2021.3064953","title":"Uncovering the Benefits and Challenges of Continuous Integration Practices","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Context (archaeology); Best practice; Suite; Process (computing); Software; Process management; Quality (philosophy); Knowledge management; Software development; Software development process; Data science; Software engineering; Business; Management","score_opus":0.032604655657231134,"score_gpt":0.25456562614295214,"score_spread":0.221960970485721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3134221463","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9325812,0.0012859943,0.03177114,0.01457691,0.00009070766,0.00026302788,0.00004054222,0.0001675348,0.019223047],"genre_scores_gemma":[0.9866556,0.00040414307,0.011694056,0.00037416857,0.000013641103,0.00011199703,0.000026294692,0.000031002048,0.0006891247],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.943271,0.03097667,0.0028618372,0.0046804436,0.014735676,0.0034743033],"domain_scores_gemma":[0.88889927,0.0719487,0.0099041,0.011604226,0.013820383,0.0038232862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05096928,0.0007649219,0.00052309164,0.003477715,0.00649462,0.012445914,0.0031419576,0.0025346857,0.0008353179],"category_scores_gemma":[0.09074743,0.0011813365,0.00047372153,0.0039870464,0.010576969,0.01504756,0.010229914,0.0053633596,0.00017910394],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001040099,0.00040378663,0.100493,0.00054342847,0.00007797191,0.0017364329,0.67422026,0.0011712548,0.0039509195,0.044321902,0.0022145864,0.17076242],"study_design_scores_gemma":[0.00008082496,0.0007499618,0.06627843,0.0017058548,0.00012812286,0.0016352,0.81095135,0.0102608325,0.003855061,0.047125358,0.057082944,0.00014608234],"about_ca_topic_score_codex":0.007836853,"about_ca_topic_score_gemma":0.011154512,"teacher_disagreement_score":0.05096928,"about_ca_system_score_codex":0.0092922,"about_ca_system_score_gemma":0.014356362,"threshold_uncertainty_score":0.26955456},"labels":[],"label_agreement":null},{"id":"W3134679844","doi":"10.21203/rs.3.rs-260432/v1","title":"FACER: An API Usage-based Code-example Recommender for Opportunistic Reuse","year":2021,"lang":"en","type":"preprint","venue":"Research Square","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Android (operating system); Java; Application programming interface; Code reuse; Source code; Snippet; Reuse; Cluster analysis; Information retrieval; Code (set theory); World Wide Web; Software; Data mining; Programming language; Artificial intelligence; Operating system; Set (abstract data type)","score_opus":0.26083632764946957,"score_gpt":0.4349643030481504,"score_spread":0.17412797539868086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3134679844","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36424384,0.0054145763,0.4743335,0.0014510928,0.00028233355,0.0013444177,0.009664127,0.13065644,0.012609654],"genre_scores_gemma":[0.38224262,0.0007177632,0.5896881,0.00046627902,0.000081953345,0.0004370995,0.013934387,0.0012984227,0.011133411],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99739397,0.0006213638,0.00018858022,0.0006150672,0.0010498442,0.00013112844],"domain_scores_gemma":[0.99316025,0.0027862496,0.00066162855,0.0013892184,0.0016743019,0.0003283092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018296901,0.0013157994,0.0010758974,0.0037824344,0.0005541027,0.0010901303,0.00218879,0.001245876,0.0035323447],"category_scores_gemma":[0.011913372,0.0005756217,0.0009234899,0.0016549565,0.00023293743,0.0023328864,0.0012462812,0.00097061764,0.0034119145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090883084,0.0012969336,0.062250793,0.001001288,0.00050625677,0.0007192465,0.00072397984,0.011926789,0.01661603,0.0019944443,0.080304086,0.82175136],"study_design_scores_gemma":[0.00027939738,0.0007412453,0.025079323,0.00020059511,0.000328152,0.0011818842,0.00049729366,0.89704686,0.016729798,0.002999255,0.054699074,0.00021710405],"about_ca_topic_score_codex":0.01216856,"about_ca_topic_score_gemma":0.03410593,"teacher_disagreement_score":0.01216856,"about_ca_system_score_codex":0.0005449558,"about_ca_system_score_gemma":0.001092496,"threshold_uncertainty_score":0.024195433},"labels":[],"label_agreement":null},{"id":"W3134882166","doi":"10.1145/3345629.3351449","title":"Does chronology matter in JIT defect prediction?","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Eclipse; Computer science; Code (set theory); Brier score; Sampling (signal processing); Data mining; Artificial intelligence; Set (abstract data type); Programming language","score_opus":0.010459454153344198,"score_gpt":0.25521669631118565,"score_spread":0.24475724215784145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3134882166","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77407986,0.056529377,0.12355466,0.01792543,0.0021449241,0.00014506107,0.0073404927,0.0017242448,0.016555937],"genre_scores_gemma":[0.97445154,0.006882221,0.011848157,0.0006925183,0.0019586205,0.000057908528,0.0020713143,0.0002825504,0.0017552883],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99758434,0.00061376236,0.00017806211,0.0011659043,0.00032655976,0.00013139784],"domain_scores_gemma":[0.8648671,0.098175585,0.0211459,0.006816151,0.006341608,0.0026536335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007302391,0.0011399172,0.00096271,0.0026720145,0.0008899232,0.0032860206,0.0013591477,0.0017412622,0.009504486],"category_scores_gemma":[0.06372259,0.00072852965,0.00069563044,0.004417577,0.0016001473,0.0057845204,0.0009215096,0.0018347418,0.0020363561],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001517741,0.00043812435,0.671753,0.0021251815,0.000559065,0.000466252,0.00090084027,0.030094795,0.002053281,0.029566683,0.02075257,0.23977251],"study_design_scores_gemma":[0.0003019117,0.0009869422,0.39411506,0.0019126846,0.00094723096,0.0014131558,0.0011828175,0.21940616,0.0039798375,0.33736038,0.038066853,0.00032700633],"about_ca_topic_score_codex":0.00554796,"about_ca_topic_score_gemma":0.0037171815,"teacher_disagreement_score":0.009504486,"about_ca_system_score_codex":0.0009217471,"about_ca_system_score_gemma":0.0013132554,"threshold_uncertainty_score":0.03861916},"labels":[],"label_agreement":null},{"id":"W3135708166","doi":"10.1145/3454122.3454124","title":"The SPACE of Developer Productivity","year":2021,"lang":"en","type":"article","venue":"Queue","topic":"Software Engineering Research","field":"Computer Science","cited_by":133,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Productivity; Computer science; Space (punctuation); Dimension (graph theory); Metric (unit); Work (physics); Software; Software engineering; Industrial engineering; Operations management; Engineering; Programming language; Mathematics; Operating system; Mechanical engineering","score_opus":0.015327004625376513,"score_gpt":0.2553325923800361,"score_spread":0.24000558775465958,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3135708166","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25366363,0.01667051,0.47658977,0.03918751,0.0015328173,0.00027394813,0.008582599,0.0025697509,0.20092946],"genre_scores_gemma":[0.93152326,0.0048869103,0.049028087,0.0005459932,0.0007011714,0.00037593153,0.0020024488,0.00041849096,0.010517787],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98549473,0.0071923686,0.0010586964,0.001788382,0.003323064,0.0011427691],"domain_scores_gemma":[0.95451957,0.025641086,0.0048455894,0.004664417,0.00632403,0.0040052985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009302644,0.0013792635,0.0008467548,0.01122466,0.001688044,0.014206544,0.0013704866,0.0019705521,0.011767569],"category_scores_gemma":[0.058890887,0.00052053016,0.0009912013,0.012192065,0.0058920267,0.017436765,0.0071241837,0.0021789705,0.0024215681],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021151704,0.000058002508,0.016235767,0.0003558242,0.00009335111,0.00017306146,0.0045544845,0.006462621,0.00064579013,0.8614504,0.01086384,0.098895356],"study_design_scores_gemma":[0.000049843216,0.00016290383,0.009925572,0.00023452245,0.000032671123,0.0002952788,0.003638002,0.011665741,0.00044957103,0.9129176,0.06054374,0.00008453386],"about_ca_topic_score_codex":0.0028090777,"about_ca_topic_score_gemma":0.0009912759,"teacher_disagreement_score":0.014206544,"about_ca_system_score_codex":0.00333445,"about_ca_system_score_gemma":0.0030018033,"threshold_uncertainty_score":0.049197674},"labels":[],"label_agreement":null},{"id":"W3135755970","doi":"10.1016/j.jss.2021.110925","title":"Building and evaluating a theory of architectural technical debt in software-intensive systems","year":2021,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"University of British Columbia","keywords":"Technical debt; Grounded theory; Construct (python library); Computer science; Architectural pattern; Software engineering; Software; Management science; Software development; Engineering; Qualitative research; Software design; Sociology; Programming language","score_opus":0.027648441915064745,"score_gpt":0.30230903048932406,"score_spread":0.2746605885742593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3135755970","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45372733,0.0016276063,0.403162,0.013565929,0.00015004726,0.0019963,0.00020914171,0.0001550823,0.12540661],"genre_scores_gemma":[0.92509276,0.00066552154,0.07193404,0.00046250966,0.00002125194,0.00094570377,0.0001267262,0.000020749487,0.0007306746],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97757655,0.015867284,0.0009652252,0.001197433,0.0034734677,0.0009200061],"domain_scores_gemma":[0.9077442,0.077100046,0.0045125056,0.0032087178,0.0062160874,0.0012183578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021991331,0.0011000889,0.00070601254,0.006978296,0.0038862093,0.009562248,0.00232412,0.002813576,0.0023044937],"category_scores_gemma":[0.045407757,0.0007440084,0.0010490032,0.0053867083,0.0195963,0.013132817,0.005612798,0.003838598,0.00027998985],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007014826,0.0005187262,0.024981184,0.0010916804,0.00010010435,0.0005148247,0.100178085,0.013553124,0.0011516124,0.7894462,0.0013720144,0.06702228],"study_design_scores_gemma":[0.000120568526,0.0005147648,0.021921638,0.0028631855,0.00024238841,0.00034867268,0.19420172,0.10018089,0.0022958026,0.65175974,0.025426602,0.00012401721],"about_ca_topic_score_codex":0.007647908,"about_ca_topic_score_gemma":0.0076976176,"teacher_disagreement_score":0.021991331,"about_ca_system_score_codex":0.01573872,"about_ca_system_score_gemma":0.012937956,"threshold_uncertainty_score":0.11630267},"labels":[],"label_agreement":null},{"id":"W3136174268","doi":"10.48550/arxiv.2103.08083","title":"EnHMM: On the Use of Ensemble HMMs and Stack Traces to Predict the Reassignment of Bug Report Fields","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Concordia University","funders":"","keywords":"Computer science; Eclipse; Hidden Markov model; Support vector machine; Stack (abstract data type); Precision and recall; Artificial intelligence; Measure (data warehouse); Field (mathematics); Machine learning; Data mining; Function (biology); Mathematics","score_opus":0.13374498453397266,"score_gpt":0.20894472533784217,"score_spread":0.07519974080386951,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3136174268","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21280798,0.004483253,0.7236095,0.0012587926,0.00062729686,0.0003507009,0.006365128,0.047775965,0.0027214526],"genre_scores_gemma":[0.73180395,0.0011041949,0.24284458,0.0008299892,0.00026563377,0.00024129394,0.016055562,0.0009011051,0.005953652],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985153,0.00037629966,0.00012733963,0.00051644875,0.00029594335,0.00016860318],"domain_scores_gemma":[0.9944074,0.0029123346,0.00046707556,0.00095326855,0.0010006065,0.00025930905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026102506,0.0022825415,0.0014329426,0.0031875777,0.00059688545,0.00091853744,0.0022725272,0.0014166132,0.0010380677],"category_scores_gemma":[0.009495171,0.0007832402,0.0013858237,0.001609207,0.0004419075,0.0020175376,0.0016067429,0.00259577,0.0012684828],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076197623,0.00064106425,0.076953426,0.00032820422,0.0006136357,0.00037926386,0.00041603763,0.26365235,0.006859507,0.0016299518,0.023329008,0.6244355],"study_design_scores_gemma":[0.000025469317,0.000112861344,0.0041214344,0.000030477519,0.000078034645,0.00007723873,0.000040011124,0.9895298,0.002354477,0.0017499849,0.0018507821,0.000029399807],"about_ca_topic_score_codex":0.030503022,"about_ca_topic_score_gemma":0.04466866,"teacher_disagreement_score":0.030503022,"about_ca_system_score_codex":0.0009340126,"about_ca_system_score_gemma":0.0019657773,"threshold_uncertainty_score":0.060650945},"labels":[],"label_agreement":null},{"id":"W3136840659","doi":"10.1109/msr52588.2021.00063","title":"Applying CodeBERT for Automated Program Repair of Java Simple Bugs","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Debugging; Java; Software bug; Source code; Automation; Software engineering; Programming language; Code (set theory); Source lines of code; Transformer; Software; Engineering; Set (abstract data type)","score_opus":0.033441104472501565,"score_gpt":0.3383997419766367,"score_spread":0.3049586375041351,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3136840659","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3288312,0.0022464155,0.5600649,0.00076730084,0.00031552673,0.00017804917,0.0028707192,0.10131695,0.003408976],"genre_scores_gemma":[0.8061492,0.00047062564,0.17897074,0.00033630972,0.00004374771,0.00010928968,0.0071233106,0.0012557697,0.005541062],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99960476,0.00006862337,0.000017970564,0.00016862931,0.000091528236,0.000048500813],"domain_scores_gemma":[0.99868816,0.00057443156,0.00012020082,0.00029562417,0.00027086202,0.000050740273],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007792395,0.0010606125,0.0004391954,0.001166679,0.00028766805,0.00048790127,0.0014081495,0.0011140539,0.0015547978],"category_scores_gemma":[0.0038978837,0.00041877493,0.0006208417,0.00055775675,0.00042228136,0.0012542563,0.00078858674,0.0012683251,0.0009589028],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005047809,0.00034628066,0.02359019,0.00042437835,0.00020223466,0.00044829655,0.00030645868,0.25697675,0.029265922,0.0025513358,0.024633,0.6607503],"study_design_scores_gemma":[0.000019402702,0.00008998221,0.001752177,0.000021332187,0.000026756916,0.000100524296,0.000030569452,0.9828918,0.01057967,0.002167792,0.0023065442,0.000013409964],"about_ca_topic_score_codex":0.010324187,"about_ca_topic_score_gemma":0.021067925,"teacher_disagreement_score":0.010324187,"about_ca_system_score_codex":0.0008583365,"about_ca_system_score_gemma":0.0013440307,"threshold_uncertainty_score":0.020528197},"labels":[],"label_agreement":null},{"id":"W3137108338","doi":"10.14722/ndss.2021.23112","title":"XDA: Accurate, Robust Disassembly with Transfer Learning","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Office of Naval Research; Amazon Web Services; National Science Foundation","keywords":"Computer science; Byte; Task (project management); Compiler; x86; Programming language; Code (set theory); Artificial intelligence; Heuristics; Function (biology); Parallel computing; Operating system; Software","score_opus":0.023688765004418084,"score_gpt":0.24708187989248212,"score_spread":0.22339311488806404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3137108338","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013778818,0.00037452724,0.9368623,0.00021080281,0.000148184,0.0001067372,0.0003091832,0.044897564,0.0033118648],"genre_scores_gemma":[0.25373167,0.00022974695,0.7253467,0.000567891,0.00008281127,0.0002806345,0.0022222386,0.002838253,0.01470018],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990632,0.00011351397,0.000048233356,0.00031468482,0.00035381978,0.0001065315],"domain_scores_gemma":[0.998387,0.00041804323,0.00014436728,0.00065061403,0.00032401824,0.00007602457],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012401054,0.0014088966,0.0009878959,0.00086811086,0.0006705096,0.0012122166,0.003229387,0.0013701048,0.008625581],"category_scores_gemma":[0.0034379978,0.00081058545,0.0007295795,0.00066471554,0.0010721249,0.0021257217,0.0024518045,0.002527656,0.006430216],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049501064,0.00021286334,0.001643112,0.00018915025,0.000083740146,0.00017008114,0.00007917915,0.24393575,0.019776203,0.006412292,0.027063165,0.69993955],"study_design_scores_gemma":[0.000024524981,0.00006389642,0.00023523862,0.000009703033,0.0000075138214,0.00006739448,0.000016771079,0.9780793,0.011033707,0.006299255,0.0041473308,0.000015328664],"about_ca_topic_score_codex":0.004003351,"about_ca_topic_score_gemma":0.005872995,"teacher_disagreement_score":0.008625581,"about_ca_system_score_codex":0.0010723941,"about_ca_system_score_gemma":0.0019784905,"threshold_uncertainty_score":0.028855383},"labels":[],"label_agreement":null},{"id":"W3137814148","doi":"10.1109/saner50967.2021.00060","title":"Anti-patterns in Modern Code Review: Symptoms and Prevalence","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Australian Research Council; Natural Sciences and Engineering Research Council of Canada","keywords":"Code review; Codebase; Computer science; Code (set theory); Software engineering; Process (computing); Code smell; Software quality; Software; Data science; Task (project management); Software development; Programming language; Engineering; Systems engineering","score_opus":0.016015866417325452,"score_gpt":0.27907058965140114,"score_spread":0.2630547232340757,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3137814148","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9852055,0.0022970983,0.0057433997,0.0020328884,0.00006424067,0.00017434898,0.00041760536,0.00031591154,0.0037491212],"genre_scores_gemma":[0.99455476,0.0008504925,0.0032170634,0.00034120245,0.000052224936,0.000090913636,0.0002789547,0.000070680966,0.0005438119],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.96876836,0.008819515,0.0050271493,0.0039242296,0.012281273,0.0011794564],"domain_scores_gemma":[0.66647834,0.13754492,0.14411077,0.014242603,0.032072134,0.005551343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009445498,0.0004272036,0.0006570353,0.009647396,0.0020067128,0.002356542,0.0013214482,0.0012487504,0.0011290874],"category_scores_gemma":[0.13666168,0.0005541943,0.0005000281,0.0076449914,0.0026753652,0.0031086432,0.0033786509,0.0014404249,0.00033637747],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023446123,0.0001626662,0.85365516,0.0009935742,0.00015250182,0.0012370503,0.022829367,0.00033222488,0.0038719159,0.0015845216,0.0039282003,0.11101833],"study_design_scores_gemma":[0.00002274246,0.00033960596,0.92588115,0.0009583667,0.00015466569,0.012278652,0.034188114,0.0032375602,0.002948947,0.0032196403,0.016640209,0.00013033266],"about_ca_topic_score_codex":0.0036283745,"about_ca_topic_score_gemma":0.006641704,"teacher_disagreement_score":0.009647396,"about_ca_system_score_codex":0.0017307155,"about_ca_system_score_gemma":0.0021765966,"threshold_uncertainty_score":0.049953222},"labels":[],"label_agreement":null},{"id":"W3138384652","doi":"10.1109/msr52588.2021.00035","title":"Leveraging Models to Reduce Test Cases in Software Repositories","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Test suite; Reduction (mathematics); Computer science; Fuzz testing; Test (biology); Test case; Compiler; Code coverage; Software; Reliability engineering; Data mining; Programming language; Machine learning; Mathematics; Engineering","score_opus":0.061812972797445964,"score_gpt":0.29976911081257457,"score_spread":0.2379561380151286,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3138384652","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12932405,0.0005331831,0.84650594,0.0009095424,0.000050656134,0.00039139154,0.00044164024,0.019769795,0.002073701],"genre_scores_gemma":[0.5196378,0.00036109347,0.47297025,0.00035074426,0.000039761384,0.0005152137,0.0023580487,0.0022348224,0.0015322773],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9927972,0.0028730836,0.0003825663,0.000905743,0.0026417852,0.00039965616],"domain_scores_gemma":[0.9776145,0.014111937,0.0016288321,0.004936994,0.0014895684,0.00021806255],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038256103,0.0021194678,0.0013627977,0.0038329554,0.00056652486,0.0024885884,0.002954406,0.0016714304,0.0017322215],"category_scores_gemma":[0.033681385,0.0015964775,0.0026633397,0.0017133541,0.0013733787,0.003933395,0.00231762,0.0023204375,0.0009628004],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039086473,0.00053486536,0.012754692,0.00047292022,0.00028121,0.00038078002,0.00044714418,0.6731313,0.029754799,0.010241658,0.0039987387,0.2676111],"study_design_scores_gemma":[0.00003369616,0.00011099132,0.0005795266,0.000028384973,0.000068151654,0.00009187132,0.00003443739,0.9772298,0.011836235,0.008538111,0.0014245161,0.000024222987],"about_ca_topic_score_codex":0.007681364,"about_ca_topic_score_gemma":0.0115625905,"teacher_disagreement_score":0.007681364,"about_ca_system_score_codex":0.0017817077,"about_ca_system_score_gemma":0.0031920366,"threshold_uncertainty_score":0.020231962},"labels":[],"label_agreement":null},{"id":"W3138741972","doi":"","title":"Proceedings of the Third International Workshop on Managing Technical Debt","year":2012,"lang":"en","type":"article","venue":"International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Debt; Computer science; Software; Engineering management; Dimension (graph theory); Software development; Engineering; Software engineering; Business; Finance","score_opus":0.03016033959334422,"score_gpt":0.2843532860365948,"score_spread":0.2541929464432506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3138741972","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007996907,0.05785759,0.118819594,0.12824725,0.18003294,0.0008893044,0.0023944322,0.0044473344,0.49931467],"genre_scores_gemma":[0.03675055,0.024757557,0.046397008,0.012745574,0.021746306,0.00084309734,0.0065080305,0.0036143486,0.8466374],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99497074,0.0014558553,0.00033503518,0.0008256954,0.0017510495,0.00066167384],"domain_scores_gemma":[0.99076337,0.0020650362,0.00034285538,0.0010969253,0.0031920325,0.0025397725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006975485,0.0014304025,0.0011223834,0.0017278658,0.0026122164,0.010296671,0.0031380216,0.0036690752,0.12761505],"category_scores_gemma":[0.012313009,0.000718076,0.0013799525,0.0017696578,0.0016238866,0.0099836355,0.006658651,0.006684931,0.057721734],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012218453,0.0001204297,0.00023333354,0.00031417087,0.000021996499,0.00015532518,0.0007469333,0.0004147934,0.0013944281,0.012703837,0.8431856,0.14058702],"study_design_scores_gemma":[0.000010436134,0.000025109826,0.00021734963,0.00018843132,0.000009918923,0.0000863934,0.00035136993,0.00025983792,0.0002996003,0.003975131,0.9945614,0.000014998939],"about_ca_topic_score_codex":0.0028415164,"about_ca_topic_score_gemma":0.005539392,"teacher_disagreement_score":0.12761505,"about_ca_system_score_codex":0.0024763837,"about_ca_system_score_gemma":0.0038606322,"threshold_uncertainty_score":0.42691487},"labels":[],"label_agreement":null},{"id":"W3139092713","doi":"10.1109/icpc52881.2021.00049","title":"API2Com: On the Improvement of Automatically Generated Code Comments Using API Documentations","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; Kelowna General Hospital; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Java; Application programming interface; Machine translation; Programming language; Syntax; Code (set theory); Software; Source code; Artificial intelligence; Context (archaeology); Transformer; Software engineering; Data mining; Machine learning","score_opus":0.0629132405425766,"score_gpt":0.3384337251924377,"score_spread":0.2755204846498611,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3139092713","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27664328,0.010007085,0.4709468,0.0027326813,0.0018648041,0.0015714723,0.020160379,0.20216455,0.01390893],"genre_scores_gemma":[0.403758,0.0022883138,0.49530315,0.0016840087,0.00038388086,0.0012182753,0.07588599,0.00507523,0.014403012],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9957099,0.0017920135,0.00024046739,0.0011836361,0.0008397308,0.00023414214],"domain_scores_gemma":[0.9849262,0.008316022,0.0007128459,0.0018175672,0.0037785629,0.00044879672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051580155,0.0028960963,0.00094761717,0.0046412996,0.00083245535,0.0014657953,0.0027770505,0.0021843547,0.0033511112],"category_scores_gemma":[0.02333229,0.00051494164,0.0014641776,0.0020992483,0.0006943409,0.0032576444,0.0017978494,0.0023777618,0.0051481104],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014929707,0.0012545981,0.014616772,0.0019241429,0.00024832282,0.00071882963,0.00084490806,0.051217534,0.014972068,0.001967649,0.09200716,0.81873506],"study_design_scores_gemma":[0.00026976023,0.00044215642,0.0045542107,0.00015708737,0.00011934348,0.0002641682,0.00026695585,0.95403254,0.014494813,0.0026184611,0.022708341,0.00007213965],"about_ca_topic_score_codex":0.024570579,"about_ca_topic_score_gemma":0.035022933,"teacher_disagreement_score":0.024570579,"about_ca_system_score_codex":0.0019147734,"about_ca_system_score_gemma":0.0032079876,"threshold_uncertainty_score":0.048855126},"labels":[],"label_agreement":null},{"id":"W3139366340","doi":"10.48550/arxiv.2103.11894","title":"Mea culpa: How developers fix their own simple bugs differently from other developers","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Commit; Java; Computer science; Scope (computer science); Software bug; Code (set theory); Simple (philosophy); Statement (logic); Software engineering; World Wide Web; Software; Database; Programming language","score_opus":0.08252759989557754,"score_gpt":0.18761306611545178,"score_spread":0.10508546621987425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3139366340","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9730397,0.0019788,0.007393146,0.000528291,0.00011402682,0.00008106977,0.011224541,0.0023168328,0.0033236023],"genre_scores_gemma":[0.9466857,0.0005095414,0.012423628,0.0002528528,0.00010320458,0.00012151591,0.036077205,0.0008370528,0.0029893967],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9958651,0.0013577334,0.00033425842,0.0012818103,0.0008799174,0.00028112507],"domain_scores_gemma":[0.96363604,0.018473836,0.007824941,0.0057944697,0.003154522,0.001116212],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0040901126,0.00074105023,0.0006008195,0.004097148,0.0007680843,0.0018883824,0.0010031837,0.00089001347,0.001310931],"category_scores_gemma":[0.031109605,0.0003271514,0.0005971333,0.0035633775,0.0005829369,0.002647573,0.0014658909,0.00093825324,0.0012379735],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004376489,0.0001822083,0.9023064,0.00042567594,0.00036863697,0.0002690959,0.0022111796,0.0020523726,0.002816071,0.00080596656,0.017752495,0.07037221],"study_design_scores_gemma":[0.00008068837,0.00039196116,0.9157905,0.00014593147,0.0002515302,0.0015033146,0.0028400517,0.03108419,0.004491661,0.0032365646,0.040058352,0.00012523483],"about_ca_topic_score_codex":0.006260417,"about_ca_topic_score_gemma":0.013320089,"teacher_disagreement_score":0.99590987,"about_ca_system_score_codex":0.0006508764,"about_ca_system_score_gemma":0.0006519526,"threshold_uncertainty_score":0.021630883},"labels":[],"label_agreement":null},{"id":"W3139836311","doi":"10.1109/iwsc.2012.6227878","title":"We have all of the clones, now what? Toward integrating clone analysis into software quality assessment","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"clone (Java method); Consistency (knowledge bases); Cloning (programming); Software; Software maintenance; Computer science; Software development; Software quality; Quality (philosophy); Software evolution; Software engineering; Data science; Software construction; Artificial intelligence; Biology; Programming language","score_opus":0.058823122854855164,"score_gpt":0.36245017771370147,"score_spread":0.30362705485884633,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3139836311","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058354314,0.0064886897,0.8811462,0.039777923,0.0005658267,0.00033771386,0.0001803195,0.0022261043,0.010922875],"genre_scores_gemma":[0.18085714,0.0037679356,0.80734646,0.003926063,0.0002474849,0.00016967283,0.00017561179,0.00039187595,0.0031177504],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9812135,0.0064226254,0.001299027,0.002395888,0.008013132,0.0006558078],"domain_scores_gemma":[0.9082096,0.030550092,0.012555626,0.0133908475,0.032136925,0.003156907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026181784,0.0014835615,0.0015059751,0.009792708,0.0031151224,0.013287434,0.002718422,0.004279891,0.0024202103],"category_scores_gemma":[0.09006106,0.0011277152,0.0012489639,0.0058317445,0.009412531,0.03077405,0.0057261377,0.0063611497,0.0012861424],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016725594,0.0003782038,0.07946879,0.0011384463,0.0002174958,0.00049149554,0.014171757,0.0036305315,0.010153825,0.21496563,0.01269514,0.66252136],"study_design_scores_gemma":[0.00008333985,0.0008514372,0.04683864,0.0031792468,0.00036712468,0.00290222,0.023237182,0.051397573,0.013695376,0.7275399,0.12928614,0.00062172324],"about_ca_topic_score_codex":0.008721619,"about_ca_topic_score_gemma":0.009617361,"teacher_disagreement_score":0.026181784,"about_ca_system_score_codex":0.0026890773,"about_ca_system_score_gemma":0.0056540905,"threshold_uncertainty_score":0.13846415},"labels":[],"label_agreement":null},{"id":"W3140248860","doi":"10.1109/icse.1999.841025","title":"Investigating quality factors in object-oriented designs: an industrial case study","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Computer Research Institute of Montréal","funders":"","keywords":"Suite; Computer science; Citation; Software; Library science; Operations research; Engineering; Operating system; History; Archaeology","score_opus":0.22319281047833303,"score_gpt":0.38046966751140276,"score_spread":0.15727685703306973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3140248860","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.922745,0.0012252808,0.053816237,0.0009453431,0.000024164296,0.0011170699,0.00014550494,0.00011507402,0.019866355],"genre_scores_gemma":[0.9454407,0.000606194,0.050074954,0.000080438615,0.00001437253,0.00026505548,0.00012554566,0.000038678067,0.003353945],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9794404,0.0125239575,0.00097521296,0.00068744284,0.005873168,0.0004999241],"domain_scores_gemma":[0.84883606,0.12590712,0.0078028394,0.0056228284,0.010505422,0.0013256998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019170757,0.00051010156,0.00038217055,0.0032444673,0.0022614754,0.0023632634,0.0009177257,0.001514746,0.0027600296],"category_scores_gemma":[0.047684733,0.00044336318,0.0006295093,0.0036674559,0.0022315334,0.0021838716,0.0017561961,0.0008606817,0.00033355053],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015911112,0.0062377574,0.17910926,0.0028748354,0.00020928419,0.0096884975,0.101871535,0.02575332,0.014103917,0.039229747,0.0053304415,0.61400026],"study_design_scores_gemma":[0.001799386,0.018889727,0.32848734,0.003186659,0.00097001856,0.01042371,0.2259572,0.11337733,0.04450325,0.07936919,0.17255215,0.00048407467],"about_ca_topic_score_codex":0.0043907063,"about_ca_topic_score_gemma":0.006757742,"teacher_disagreement_score":0.019170757,"about_ca_system_score_codex":0.0037229413,"about_ca_system_score_gemma":0.0026437517,"threshold_uncertainty_score":0.10138589},"labels":[],"label_agreement":null},{"id":"W3140924550","doi":"10.1109/iwsc.2012.6227863","title":"Dispersion of changes in cloned and non-cloned code","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Java; Computer science; Cloning (programming); Software maintenance; clone (Java method); Asynchronous communication; Source lines of code; Software evolution; Software; Subject (documents); Source code; Biology; Programming language; Software system; Parallel computing; Genetics; Gene; World Wide Web; Telecommunications","score_opus":0.01828286285838074,"score_gpt":0.26593363069498055,"score_spread":0.2476507678365998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3140924550","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9665799,0.0004623906,0.030233258,0.000064472515,0.000030156194,0.000067511864,0.00044315096,0.00096390606,0.0011551528],"genre_scores_gemma":[0.9790229,0.00012931644,0.019065829,0.000022567461,0.000012116527,0.00005642616,0.00079983164,0.00016281797,0.0007281508],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9909587,0.0010943452,0.00075986417,0.0018050793,0.005062475,0.0003194844],"domain_scores_gemma":[0.9343045,0.031578373,0.011068749,0.010030078,0.012028852,0.0009895696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038210135,0.00036932444,0.0004532453,0.004147248,0.0004765922,0.0010533761,0.00074967643,0.0004746642,0.0004973832],"category_scores_gemma":[0.039604437,0.0003455741,0.0003730774,0.0029259722,0.0006974908,0.0017857986,0.0009067363,0.00083865726,0.00016613746],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016188944,0.0004963608,0.37205213,0.0006732608,0.00044355806,0.0008975993,0.0051932367,0.03410129,0.2072579,0.0029937087,0.0011533154,0.37311876],"study_design_scores_gemma":[0.00004823128,0.0013465057,0.5854566,0.00011096945,0.00030896394,0.0014662686,0.0012851822,0.17599738,0.22463228,0.0038564128,0.005332549,0.00015875702],"about_ca_topic_score_codex":0.0020141522,"about_ca_topic_score_gemma":0.0025152136,"teacher_disagreement_score":0.004147248,"about_ca_system_score_codex":0.00071741967,"about_ca_system_score_gemma":0.00042935024,"threshold_uncertainty_score":0.020207703},"labels":[],"label_agreement":null},{"id":"W3141218678","doi":"10.1109/msr.2012.6224300","title":"Think locally, act globally: Improving defect and effort prediction models","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Predictive modelling; Software; Machine learning; Contrast (vision); Data mining; Data science; Data modeling; Multitude; Artificial intelligence; Software engineering","score_opus":0.014965806582333154,"score_gpt":0.23888365720614432,"score_spread":0.22391785062381117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3141218678","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26105765,0.0018603543,0.7265628,0.0024976558,0.00012230212,0.000116876254,0.00079164415,0.0031610911,0.0038296129],"genre_scores_gemma":[0.8734928,0.0005836987,0.120872855,0.000527091,0.00017958201,0.0001487171,0.0015277817,0.00034420672,0.00232335],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975211,0.0010885025,0.00011102603,0.00080110756,0.00030910937,0.00016913858],"domain_scores_gemma":[0.9884567,0.008003178,0.0010114431,0.001180785,0.0010175646,0.000330249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074155577,0.0027798673,0.0022458376,0.0024162487,0.00059866766,0.0020413073,0.0020362341,0.001742364,0.0009766965],"category_scores_gemma":[0.016243823,0.0008362198,0.0015457062,0.002172394,0.00075897307,0.0049420567,0.0018466857,0.0030004878,0.00072852895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033390924,0.0004188058,0.07665474,0.0001158509,0.0005525135,0.00019104502,0.00054235687,0.70096934,0.0014174336,0.003389546,0.0047136825,0.21070091],"study_design_scores_gemma":[0.000010628813,0.00007093898,0.0030386893,0.000014579078,0.0000398824,0.000023897706,0.000058312202,0.99230736,0.00025901818,0.0038556666,0.0003063332,0.000014781885],"about_ca_topic_score_codex":0.009455836,"about_ca_topic_score_gemma":0.012553582,"teacher_disagreement_score":0.009455836,"about_ca_system_score_codex":0.000777506,"about_ca_system_score_gemma":0.0011468268,"threshold_uncertainty_score":0.03921771},"labels":[],"label_agreement":null},{"id":"W3142784638","doi":"10.1109/asonam49781.2020.9381373","title":"Unraveling the Semantic Evolution of Core Nodes in a Global Contribution Network","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Ontology; Task (project management); Context (archaeology); Semantic network; Data science; Domain (mathematical analysis); Core (optical fiber); Domain knowledge; Social network (sociolinguistics); Artificial intelligence; World Wide Web; Social media","score_opus":0.02121762645643839,"score_gpt":0.26453811467939087,"score_spread":0.2433204882229525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3142784638","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6803227,0.00042687153,0.30478373,0.0009319268,0.000036975303,0.00013160768,0.0009457706,0.00044309668,0.0119773215],"genre_scores_gemma":[0.9303597,0.00021015019,0.06698787,0.0000358987,0.000011063259,0.000054641296,0.0005926324,0.00008688916,0.0016610549],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99924624,0.00028526556,0.000030969404,0.00022421882,0.00015560821,0.00005778435],"domain_scores_gemma":[0.99431473,0.0033506374,0.00089123443,0.00052109984,0.0006538188,0.00026839375],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0015765362,0.00034013932,0.0002834736,0.0031964278,0.0008743085,0.0019442885,0.00048668848,0.00065462146,0.0013545016],"category_scores_gemma":[0.008990003,0.00024924846,0.00030621036,0.0026864468,0.0010096321,0.004967295,0.001330425,0.0006942502,0.0002240995],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075465796,0.00029891895,0.28243363,0.00058324117,0.00024544707,0.0015799802,0.022803348,0.093530096,0.043233152,0.20961183,0.0046359147,0.34028974],"study_design_scores_gemma":[0.00003372658,0.00015598998,0.12656115,0.00014919772,0.00025658894,0.00083513703,0.011043682,0.6271078,0.013600708,0.1877635,0.03240473,0.00008770045],"about_ca_topic_score_codex":0.0073756143,"about_ca_topic_score_gemma":0.013817108,"teacher_disagreement_score":0.9968036,"about_ca_system_score_codex":0.0011609138,"about_ca_system_score_gemma":0.0008954881,"threshold_uncertainty_score":0.014665365},"labels":[],"label_agreement":null},{"id":"W3143486204","doi":"10.1109/icse.2012.6227036","title":"On the analysis of evolution of software artefacts and programs","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Macro; Computer science; Software evolution; Construct (python library); Tree (set theory); Field (mathematics); Data science; Software; Artificial intelligence; Theoretical computer science; Software development; Programming language","score_opus":0.021472892918729707,"score_gpt":0.254725788877641,"score_spread":0.23325289595891133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3143486204","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05467731,0.006374891,0.9316552,0.0006194571,0.000051858366,0.00018144265,0.00073215103,0.0010687911,0.0046388735],"genre_scores_gemma":[0.2479208,0.0046400144,0.74175805,0.00017562156,0.00010334533,0.00031155845,0.0017883953,0.00036705515,0.0029351688],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9959448,0.0014208644,0.0003193884,0.00087857066,0.0012338121,0.00020259937],"domain_scores_gemma":[0.9841575,0.01075241,0.0021136778,0.0015096717,0.0012365655,0.00023025331],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029273415,0.0011198946,0.0009802856,0.012293559,0.0011656904,0.0033067006,0.0012021902,0.0012149584,0.0016119527],"category_scores_gemma":[0.017483246,0.00055882044,0.0018071986,0.009494068,0.0033968643,0.0038027924,0.0016099871,0.0012665457,0.00044196693],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025600925,0.00019998912,0.06287208,0.002391911,0.00039065984,0.0013842263,0.0049591856,0.10212033,0.019924294,0.23467112,0.0021243088,0.56870586],"study_design_scores_gemma":[0.000040069826,0.00028848898,0.07226763,0.0011900415,0.00030391917,0.0023900159,0.0019336316,0.41653717,0.0209467,0.4100771,0.07379418,0.00023103211],"about_ca_topic_score_codex":0.005762847,"about_ca_topic_score_gemma":0.0033609567,"teacher_disagreement_score":0.012293559,"about_ca_system_score_codex":0.0018876529,"about_ca_system_score_gemma":0.0014371189,"threshold_uncertainty_score":0.015481472},"labels":[],"label_agreement":null},{"id":"W3143876921","doi":"10.1007/978-3-030-73128-1_5","title":"Automatically Classifying Non-functional Requirements with Feature Extraction and Supervised Machine Learning Techniques: A Research Preview","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Support vector machine; Artificial intelligence; Machine learning; Context (archaeology); Feature extraction; Precision and recall; Process (computing); Decision tree; Data mining; Pattern recognition (psychology)","score_opus":0.05603098863808203,"score_gpt":0.32850835352302515,"score_spread":0.2724773648849431,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3143876921","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015320161,0.017100364,0.9483141,0.000830215,0.00028756933,0.00016851486,0.0004991897,0.0036709316,0.013808885],"genre_scores_gemma":[0.122843094,0.02201124,0.80557245,0.00039800777,0.0005463802,0.00014682942,0.0041413936,0.0006635734,0.04367711],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99927,0.00007653196,0.000056980563,0.00018308323,0.0003742389,0.000039156832],"domain_scores_gemma":[0.99870145,0.0005680029,0.000103805825,0.00016977289,0.00042240776,0.000034586945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010237708,0.0007845736,0.00069980155,0.0026699875,0.00024569698,0.001706369,0.0013602853,0.0008401573,0.004042679],"category_scores_gemma":[0.0018769437,0.00049584825,0.0011997038,0.0023434989,0.00043167267,0.0021822264,0.00037355482,0.00086461066,0.004039044],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035605495,0.00012348668,0.0011963433,0.00047742235,0.000026710299,0.00008324983,0.00008680323,0.0037024727,0.023798732,0.005347782,0.009195885,0.95592546],"study_design_scores_gemma":[0.00003672638,0.0007460115,0.020242589,0.0014596964,0.00020567466,0.0042134323,0.00062123354,0.41470614,0.14908434,0.07082,0.337618,0.0002461972],"about_ca_topic_score_codex":0.0013871818,"about_ca_topic_score_gemma":0.0019842912,"teacher_disagreement_score":0.004042679,"about_ca_system_score_codex":0.0004692448,"about_ca_system_score_gemma":0.0005675077,"threshold_uncertainty_score":0.013524115},"labels":[],"label_agreement":null},{"id":"W3143959408","doi":"","title":"Software library for reuse-oriented program development.","year":2000,"lang":"en","type":"article","venue":"Scholarship at UWindsor (University of Windsor)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Windsor","funders":"","keywords":"Reuse; Software engineering; Computer science; Software development; Software; Package development process; Programming language; Software construction; Engineering","score_opus":0.01766324495393714,"score_gpt":0.23113983746142178,"score_spread":0.21347659250748463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3143959408","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010283993,0.0022700182,0.74924153,0.0011851838,0.00059566373,0.00091758615,0.0075649917,0.13360539,0.10359119],"genre_scores_gemma":[0.014165893,0.0054731565,0.809818,0.0015663399,0.0005974723,0.0024793083,0.037531484,0.02453573,0.1038326],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99713373,0.0005061395,0.0005211608,0.00037509174,0.0012514427,0.00021248621],"domain_scores_gemma":[0.9946616,0.0017328106,0.00039892993,0.001620463,0.0011776495,0.0004085909],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003073781,0.0017534574,0.0015075397,0.0044271,0.0009895397,0.0032093602,0.0030083163,0.0018108013,0.097839236],"category_scores_gemma":[0.0095564015,0.0015345359,0.0019161645,0.0050353794,0.0013359549,0.0058993543,0.004623203,0.004661798,0.116322175],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021131657,0.00023083006,0.00074558496,0.002191529,0.00014042419,0.00037997388,0.0004877189,0.0024766708,0.007952072,0.12553394,0.31850538,0.5411446],"study_design_scores_gemma":[0.00015726955,0.000102871214,0.00067130115,0.00052076025,0.00006592874,0.00063748437,0.000055312572,0.0065184804,0.00543057,0.05573546,0.93003184,0.00007280727],"about_ca_topic_score_codex":0.0020386337,"about_ca_topic_score_gemma":0.0020321375,"teacher_disagreement_score":0.097839236,"about_ca_system_score_codex":0.0010778373,"about_ca_system_score_gemma":0.0035378663,"threshold_uncertainty_score":0.32730484},"labels":[],"label_agreement":null},{"id":"W3144106047","doi":"10.1109/msr.2009.5069475","title":"The promises and perils of mining git","year":2009,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":310,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Commit; Computer science; Decentralization; Semantics (computer science); Confusion; Focus (optics); Data science; Code (set theory); Source code; World Wide Web; Computer security; Database; Political science; Programming language; Set (abstract data type)","score_opus":0.014940648768104551,"score_gpt":0.26014146477459754,"score_spread":0.245200816006493,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3144106047","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3025976,0.025483241,0.4628119,0.17616287,0.002008006,0.000865939,0.011125077,0.0051580463,0.013787389],"genre_scores_gemma":[0.42660075,0.0054477206,0.544447,0.0056613744,0.0023673908,0.001037268,0.010824833,0.0013454894,0.0022682191],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.854779,0.08575933,0.00889255,0.011312301,0.036832146,0.0024247738],"domain_scores_gemma":[0.38585538,0.43123236,0.027273024,0.11434689,0.03490128,0.00639111],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.16516906,0.0018022663,0.002848276,0.021045513,0.0059562167,0.01947621,0.007215227,0.006229762,0.0011402869],"category_scores_gemma":[0.45309782,0.0020389103,0.0019025964,0.022348382,0.01528206,0.035941035,0.012683919,0.012944236,0.0013354123],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011775953,0.0005846863,0.22134125,0.004341671,0.0007659025,0.0016976433,0.02163153,0.027966259,0.004444677,0.13997601,0.053772263,0.5223005],"study_design_scores_gemma":[0.00019782675,0.00027018692,0.056807246,0.0020255682,0.00020752639,0.0021876604,0.0208345,0.13038836,0.005486655,0.6357783,0.14527722,0.00053892867],"about_ca_topic_score_codex":0.005088204,"about_ca_topic_score_gemma":0.008103662,"teacher_disagreement_score":0.83483094,"about_ca_system_score_codex":0.0047747553,"about_ca_system_score_gemma":0.008772228,"threshold_uncertainty_score":0.8735079},"labels":[],"label_agreement":null},{"id":"W3144273096","doi":"10.1109/tse.2021.3068901","title":"On the Untriviality of Trivial Packages: An Empirical Study of npm JavaScript Packages","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Concordia University","funders":"","keywords":"Computer science; JavaScript; Programming language; Empirical research; Software engineering","score_opus":0.0341506444947318,"score_gpt":0.29139461396325345,"score_spread":0.2572439694685216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3144273096","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9963972,0.00018832619,0.0014847482,0.00018046118,0.0000056578815,0.000019920262,0.00019654445,0.00002873164,0.0014983782],"genre_scores_gemma":[0.9971624,0.0001347639,0.001618977,0.00007351218,0.000014850143,0.000032662138,0.00047175188,0.00006156699,0.00042946637],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9927089,0.0034006615,0.00041724424,0.0012520032,0.0017496251,0.00047151954],"domain_scores_gemma":[0.86486655,0.09856489,0.02085099,0.0050157434,0.0071316822,0.0035701983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0087051485,0.00035556045,0.00044865717,0.003949329,0.0014463059,0.0022348643,0.001244349,0.0010294714,0.0019779492],"category_scores_gemma":[0.06867024,0.00038785604,0.0003605428,0.0040819743,0.00270992,0.006343861,0.0021090866,0.0016170955,0.000784976],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013891289,0.00024321358,0.943463,0.00030202328,0.00008365683,0.0005277758,0.023863818,0.00052909896,0.0012270822,0.0027136344,0.002616284,0.02429156],"study_design_scores_gemma":[0.000018344417,0.0002153545,0.9292747,0.00021874021,0.00006232598,0.001490836,0.042642456,0.011149407,0.0009153609,0.002690806,0.01126432,0.00005733596],"about_ca_topic_score_codex":0.0022061467,"about_ca_topic_score_gemma":0.003401791,"teacher_disagreement_score":0.0087051485,"about_ca_system_score_codex":0.0006191703,"about_ca_system_score_gemma":0.00051761576,"threshold_uncertainty_score":0.046037793},"labels":[],"label_agreement":null},{"id":"W3144363571","doi":"10.1109/msr.2012.6224267","title":"Bug introducing changes: A case study with Android","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Android (operating system); Computer science; Software bug; Software maintenance; Software development; Security bug; Software engineering; Software; Computer security; Operating system; Software security assurance; Information security","score_opus":0.02477399519306672,"score_gpt":0.2782318188220735,"score_spread":0.25345782362900676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3144363571","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99510103,0.00029651105,0.002611336,0.00046622605,0.000015782258,0.00014372975,0.00013110765,0.00005415896,0.0011800796],"genre_scores_gemma":[0.9889543,0.00045879086,0.008693838,0.00018099755,0.00002194916,0.000066032204,0.00012667119,0.00005235792,0.0014450925],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99707067,0.0013546398,0.00024429616,0.00033174534,0.0007622444,0.00023642185],"domain_scores_gemma":[0.9714536,0.021998517,0.0027443548,0.0013409387,0.001619879,0.0008425907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023062867,0.00067436113,0.0004551844,0.002477318,0.002400174,0.001012114,0.0012041637,0.002853119,0.00078434596],"category_scores_gemma":[0.020898046,0.0005112133,0.00060679356,0.0016927529,0.0016238627,0.0015106522,0.0013504911,0.0014203151,0.00018638671],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006942696,0.0044490495,0.3742671,0.0013782933,0.00027332507,0.2463471,0.19404809,0.006130257,0.01875878,0.0030233937,0.005569712,0.14506063],"study_design_scores_gemma":[0.0003493964,0.004086362,0.550148,0.000754905,0.00058871735,0.18411496,0.15762396,0.027447067,0.026740512,0.004455017,0.043286704,0.0004044729],"about_ca_topic_score_codex":0.010404995,"about_ca_topic_score_gemma":0.03110791,"teacher_disagreement_score":0.010404995,"about_ca_system_score_codex":0.0009871954,"about_ca_system_score_gemma":0.0007457941,"threshold_uncertainty_score":0.020688891},"labels":[],"label_agreement":null},{"id":"W3144571850","doi":"10.1109/icse.2012.6227086","title":"CodeTimeline: Storytelling with versioning data","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Storytelling; Computer science; Software versioning; Casual; Visualization; World Wide Web; Narrative; Data visualization; Human–computer interaction; Software engineering; Software; Programming language; Artificial intelligence","score_opus":0.047263827506573264,"score_gpt":0.28291189502599123,"score_spread":0.23564806751941797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3144571850","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042856038,0.00041585873,0.8682029,0.00083816913,0.00017795296,0.0009007189,0.0035424898,0.070455626,0.0126102185],"genre_scores_gemma":[0.27056918,0.00044477853,0.7012665,0.00029746816,0.00011992478,0.0016238323,0.0063813897,0.008875387,0.010421527],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99823135,0.0009805433,0.00011505008,0.00025506923,0.00033652602,0.00008145911],"domain_scores_gemma":[0.98258704,0.013299048,0.0006310396,0.0022755968,0.00066257507,0.0005446614],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028283685,0.0015763832,0.00051214144,0.0015566583,0.0006641469,0.0024930846,0.0022282137,0.0014344922,0.01778453],"category_scores_gemma":[0.018646093,0.00071238825,0.0006411108,0.0008939962,0.00087471853,0.0046418062,0.003344,0.0015507916,0.0022999505],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017368668,0.0006910288,0.0059947856,0.0029595792,0.00022658167,0.0018226162,0.033871777,0.015242991,0.056550264,0.022777708,0.082017556,0.77610826],"study_design_scores_gemma":[0.0009065948,0.0017922473,0.008354444,0.0008595828,0.00027229037,0.0026612096,0.0059985323,0.23592974,0.09234952,0.049855717,0.60055137,0.00046884545],"about_ca_topic_score_codex":0.00082899735,"about_ca_topic_score_gemma":0.0015146803,"teacher_disagreement_score":0.01778453,"about_ca_system_score_codex":0.0003576226,"about_ca_system_score_gemma":0.0003744931,"threshold_uncertainty_score":0.05949521},"labels":[],"label_agreement":null},{"id":"W3145055218","doi":"10.1109/iwsc.2012.6227875","title":"Shuffling and randomization for scalable source code clone detection","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; University of Saskatchewan; Concordia University","funders":"","keywords":"Shuffling; Computer science; Scalability; Code (set theory); Source code; clone (Java method); Key (lock); State (computer science); Theoretical computer science; Computer engineering; Data mining; Machine learning; Programming language; Database; Operating system","score_opus":0.01647478806373406,"score_gpt":0.2557078956745717,"score_spread":0.23923310761083763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3145055218","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057288263,0.0006173046,0.89556694,0.00028059864,0.000115841016,0.00041109297,0.0007664555,0.04361826,0.0013352913],"genre_scores_gemma":[0.2768909,0.00021303182,0.71692467,0.0002572702,0.00011664232,0.00061784167,0.0023257514,0.0013891015,0.0012648009],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912411,0.0028149355,0.0007986648,0.0021359765,0.0026520558,0.0003571339],"domain_scores_gemma":[0.9697995,0.012348627,0.0026224367,0.012998102,0.0018361816,0.00039501922],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062040375,0.00138049,0.0013736992,0.0038021817,0.0009180575,0.0018272183,0.0026616962,0.0013293802,0.0020315584],"category_scores_gemma":[0.030601127,0.0006967876,0.001198408,0.0031805644,0.0012044602,0.0050308215,0.0023441592,0.0016239594,0.0017836876],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009855728,0.00043627966,0.013098986,0.00048740004,0.0002752421,0.00043918513,0.0003877762,0.028392008,0.084641136,0.014717196,0.016334413,0.83980477],"study_design_scores_gemma":[0.00034097667,0.0009781857,0.007778328,0.00011421266,0.00019938787,0.0013909665,0.00024693104,0.6963714,0.21908201,0.043131586,0.030102732,0.00026325227],"about_ca_topic_score_codex":0.0009236473,"about_ca_topic_score_gemma":0.0015168958,"teacher_disagreement_score":0.0062040375,"about_ca_system_score_codex":0.0007767851,"about_ca_system_score_gemma":0.0020186612,"threshold_uncertainty_score":0.03281051},"labels":[],"label_agreement":null},{"id":"W3145503536","doi":"10.1109/icse.2012.6227174","title":"Automated analysis of CSS rules to support style maintenance","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Cascading Style Sheets; Web application; Semantics (computer science); Inheritance (genetic algorithm); Semantic Web; Class (philosophy); World Wide Web; Software engineering; Programming language; Information retrieval; Web page; Artificial intelligence","score_opus":0.01899203707600168,"score_gpt":0.29630894706467764,"score_spread":0.277316909988676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3145503536","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17597684,0.00048057185,0.7232594,0.00028182942,0.00011098846,0.00053169497,0.003214714,0.092502646,0.0036414368],"genre_scores_gemma":[0.36602283,0.00026379628,0.6208693,0.00011261094,0.00006603819,0.00021139134,0.0067312554,0.0032297918,0.0024931245],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99665606,0.00060057855,0.00042770675,0.00068267545,0.0014558256,0.00017720557],"domain_scores_gemma":[0.9731216,0.01019175,0.0029791533,0.0062075723,0.007129486,0.000370476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022827,0.00121566,0.0009766383,0.0069201253,0.0007719767,0.0021303887,0.0018210843,0.0008356256,0.0021091518],"category_scores_gemma":[0.021070302,0.0006221745,0.00094605435,0.0022386366,0.00047844453,0.0014103169,0.0009911923,0.0009509153,0.0021828006],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036410594,0.00041401928,0.055079542,0.0005802521,0.00016803229,0.0008162231,0.0012093744,0.012886967,0.09854708,0.0024258813,0.017241744,0.8102669],"study_design_scores_gemma":[0.000109974615,0.0001758959,0.03599106,0.00013928582,0.00018421974,0.0015316862,0.00029405457,0.7119686,0.21489206,0.006190136,0.028338967,0.00018409082],"about_ca_topic_score_codex":0.003916542,"about_ca_topic_score_gemma":0.0051632444,"teacher_disagreement_score":0.0069201253,"about_ca_system_score_codex":0.0006884919,"about_ca_system_score_gemma":0.0018878351,"threshold_uncertainty_score":0.012072206},"labels":[],"label_agreement":null},{"id":"W3145602566","doi":"10.1145/3586074","title":"A Practical Survey on Faster and Lighter Transformers","year":2023,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Software Engineering Research","field":"Computer Science","cited_by":120,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; Polytechnique Montréal","funders":"","keywords":"Computer science; Transformer; Computation; Machine learning; Artificial intelligence; Algorithm","score_opus":0.18923570636933304,"score_gpt":0.41316001859904156,"score_spread":0.22392431222970852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3145602566","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021718252,0.87139285,0.09306546,0.00339079,0.0012140437,0.00010086828,0.00038688013,0.0010957349,0.027181616],"genre_scores_gemma":[0.017385785,0.8995611,0.06185056,0.0020863195,0.0016578602,0.00017108837,0.0012093729,0.0006125087,0.015465452],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986131,0.00025170116,0.0001407102,0.0002678346,0.00061898615,0.000107603475],"domain_scores_gemma":[0.99561006,0.00258131,0.00018847699,0.00044186515,0.0010716153,0.00010672446],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020583586,0.0017032858,0.0011452858,0.0035645156,0.00046010737,0.001991832,0.0026674694,0.0016154583,0.019901756],"category_scores_gemma":[0.009868125,0.0011558023,0.0010950107,0.0055422825,0.0008385274,0.006884668,0.00152784,0.002498083,0.012619],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007393943,0.00006784573,0.00033249796,0.005085119,0.000054178625,0.000059550643,0.000074759664,0.0024119702,0.0014377885,0.040829286,0.048782606,0.9007905],"study_design_scores_gemma":[0.000032539112,0.00016686729,0.00055231905,0.0031137818,0.000100828336,0.00073800486,0.00008808194,0.006716132,0.0026297276,0.031480238,0.9543206,0.000060837465],"about_ca_topic_score_codex":0.001813666,"about_ca_topic_score_gemma":0.0021989986,"teacher_disagreement_score":0.019901756,"about_ca_system_score_codex":0.00087001495,"about_ca_system_score_gemma":0.0017946971,"threshold_uncertainty_score":0.06657797},"labels":[],"label_agreement":null},{"id":"W3145867603","doi":"10.1109/mtd.2012.6226002","title":"On the role of requirements in understanding and managing technical debt","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Technical debt; Requirements analysis; Non-functional requirement; Computer science; Business requirements; Requirements management; Risk analysis (engineering); System requirements specification; Software requirements specification; Stakeholder; Requirement; Software quality; Systems engineering; Software; Software engineering; Engineering; Software system; Software development; Business; Software design; Business process; Operations management; Work in process; Software construction","score_opus":0.05384467009957516,"score_gpt":0.2914664233100002,"score_spread":0.23762175321042506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3145867603","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05488128,0.0030661088,0.8036823,0.049832866,0.0001605692,0.00017938824,0.00012261774,0.00027473547,0.0878001],"genre_scores_gemma":[0.7278646,0.0028734305,0.25919265,0.0026018792,0.0002106244,0.00029424587,0.000150216,0.00032116083,0.0064912066],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98174596,0.011397229,0.0015254515,0.0012150749,0.0030700427,0.0010463424],"domain_scores_gemma":[0.91245794,0.06578615,0.006628675,0.0070105228,0.006809265,0.001307579],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026045145,0.0010733769,0.0007581543,0.004063027,0.003324638,0.010004003,0.0022549927,0.004301166,0.0034768526],"category_scores_gemma":[0.071101286,0.0011528943,0.0009112365,0.003339293,0.018291833,0.032377355,0.005970261,0.0064691906,0.0006089958],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000024460878,0.000027500557,0.0011422044,0.00011210595,0.000012017129,0.00023660921,0.0057963068,0.0061657424,0.0003581287,0.96101624,0.001538632,0.023570003],"study_design_scores_gemma":[0.000011025372,0.000028696044,0.0009160199,0.0002778735,0.000013503018,0.00027754297,0.0042182794,0.018764103,0.00053506,0.95149195,0.023427684,0.000038236096],"about_ca_topic_score_codex":0.0059598256,"about_ca_topic_score_gemma":0.0037714215,"teacher_disagreement_score":0.026045145,"about_ca_system_score_codex":0.0056629977,"about_ca_system_score_gemma":0.0040805675,"threshold_uncertainty_score":0.1377415},"labels":[],"label_agreement":null},{"id":"W3146268108","doi":"10.1109/msr.2012.6224296","title":"A Linked Data platform for mining software repositories","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software versioning; Software; World Wide Web; Cloud computing; Software mining; BitTorrent tracker; Software engineering; Database; Software development; Data science; Software construction","score_opus":0.0885113870139768,"score_gpt":0.3235739738579704,"score_spread":0.23506258684399356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3146268108","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011380826,0.00095437735,0.584897,0.0013175859,0.00025168806,0.003188681,0.27740678,0.11056896,0.010034086],"genre_scores_gemma":[0.033798076,0.0008042714,0.46570432,0.00039729,0.0000978983,0.003202811,0.48837417,0.0036908085,0.003930292],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99226016,0.001242904,0.0016050179,0.001602502,0.0029787011,0.00031068258],"domain_scores_gemma":[0.98139036,0.004865831,0.00238151,0.0070703067,0.003116186,0.0011758967],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073424787,0.0017091231,0.0014038613,0.019794432,0.0025805237,0.005626353,0.0039228494,0.0022280007,0.0076976176],"category_scores_gemma":[0.030596783,0.0016175383,0.0024594632,0.021601945,0.000806864,0.009088364,0.00850551,0.0030064634,0.0067618494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013980023,0.0010164788,0.028401058,0.005185725,0.0014484159,0.0017029326,0.0028264185,0.03308338,0.014583257,0.12237898,0.30899116,0.47898418],"study_design_scores_gemma":[0.0004939862,0.00032935015,0.017636666,0.0010072326,0.00037362214,0.00089040725,0.0009585638,0.13826485,0.01949439,0.14927955,0.67080045,0.00047090658],"about_ca_topic_score_codex":0.010257907,"about_ca_topic_score_gemma":0.010569367,"teacher_disagreement_score":0.019794432,"about_ca_system_score_codex":0.0021224166,"about_ca_system_score_gemma":0.005028085,"threshold_uncertainty_score":0.038831174},"labels":[],"label_agreement":null},{"id":"W3146730906","doi":"10.1002/smr.2343","title":"On the value of filter feature selection techniques in homogeneous ensembles effort estimation","year":2021,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Feature selection; Preprocessor; Artificial intelligence; Data mining; Software; Filter (signal processing); Support vector machine; Feature (linguistics); Pattern recognition (psychology); Multilayer perceptron; Machine learning; Subspace topology; Random forest; k-nearest neighbors algorithm; Dimensionality reduction; Process (computing); Curse of dimensionality; Homogeneous; Artificial neural network; Mathematics","score_opus":0.008875852151698763,"score_gpt":0.26110152497721284,"score_spread":0.25222567282551406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3146730906","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29283068,0.00077104376,0.7044649,0.00025235722,0.000052045725,0.000065584725,0.00014850646,0.00046737675,0.00094756647],"genre_scores_gemma":[0.9085015,0.00015033493,0.0906736,0.00005460084,0.000038859063,0.000057324476,0.00024242901,0.000036190042,0.00024509494],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975286,0.001364898,0.00014554037,0.00038227122,0.00043297943,0.00014568299],"domain_scores_gemma":[0.9841945,0.012478599,0.00063029205,0.0009835493,0.0015804581,0.0001325576],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065481244,0.00088703167,0.00096256513,0.0017331534,0.0004685087,0.00089045026,0.0006597866,0.00088503596,0.0005996959],"category_scores_gemma":[0.022206696,0.00022059858,0.00083374535,0.0012475556,0.00039464128,0.0015266853,0.0006240755,0.0007195872,0.00014350691],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051237806,0.00030063227,0.025630757,0.00011236968,0.00050903496,0.0000834508,0.00011238929,0.6273342,0.0060068886,0.0032787567,0.0011430422,0.3349761],"study_design_scores_gemma":[0.00001195534,0.00012622606,0.0042442963,0.000014112825,0.000052278472,0.000020664718,0.00003150876,0.99119747,0.002858717,0.0012644595,0.00016744818,0.000010948825],"about_ca_topic_score_codex":0.004242222,"about_ca_topic_score_gemma":0.00291803,"teacher_disagreement_score":0.0065481244,"about_ca_system_score_codex":0.00042807843,"about_ca_system_score_gemma":0.00067936064,"threshold_uncertainty_score":0.03463018},"labels":[],"label_agreement":null},{"id":"W3146926538","doi":"10.1109/icse.2012.6227238","title":"WorkItemExplorer: Visualizing software development tasks using an interactive exploration environment","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Data exploration; Task (project management); Ask price; Human–computer interaction; Software; Software development; Task management; Visualization; Task analysis; Data visualization; Interactive visualization; Software engineering; Artificial intelligence; Systems engineering; Programming language; Engineering","score_opus":0.09802627681252768,"score_gpt":0.3253566351947563,"score_spread":0.22733035838222865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3146926538","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08827323,0.0010632045,0.75533414,0.0014013617,0.00022717458,0.00075694534,0.013488753,0.10430147,0.03515374],"genre_scores_gemma":[0.27167144,0.0012618995,0.6937789,0.0005245657,0.00008566174,0.001362971,0.010566708,0.008776197,0.011971696],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99954283,0.00015636603,0.000028534434,0.000068677255,0.00012279749,0.000080821235],"domain_scores_gemma":[0.9978321,0.0014356226,0.00008515363,0.00025927194,0.00014969149,0.00023823352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010232906,0.0014685176,0.00048106562,0.001673083,0.00057268905,0.0024190466,0.001381114,0.0010312869,0.016507944],"category_scores_gemma":[0.0031335487,0.0007004531,0.0009805397,0.00083206285,0.00049831806,0.0020975799,0.0038731936,0.0011166757,0.0025119402],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004101545,0.0009678234,0.015678898,0.0033534947,0.0004204833,0.002765979,0.029048324,0.02975715,0.10205817,0.02119837,0.24401756,0.5466321],"study_design_scores_gemma":[0.0008304441,0.0008655857,0.022217935,0.0011514109,0.00027698185,0.0021881072,0.0067094364,0.15944903,0.088698484,0.02953844,0.68720233,0.00087188673],"about_ca_topic_score_codex":0.0032957725,"about_ca_topic_score_gemma":0.0071952897,"teacher_disagreement_score":0.016507944,"about_ca_system_score_codex":0.00031148116,"about_ca_system_score_gemma":0.0006658651,"threshold_uncertainty_score":0.055224597},"labels":[],"label_agreement":null},{"id":"W3147021188","doi":"10.1109/iwsc.2012.6227873","title":"Near-miss model clone detection for Simulink models","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; General Motors of Canada","keywords":"Computer science; Leverage (statistics); clone (Java method); Granularity; Source code; Identification (biology); Detector; Artificial intelligence; Data mining; Programming language","score_opus":0.05080366454702209,"score_gpt":0.2929140153861381,"score_spread":0.242110350839116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3147021188","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07625918,0.0001438815,0.90825295,0.00013376845,0.000031258252,0.00010290821,0.00020982767,0.013935691,0.0009304473],"genre_scores_gemma":[0.45983806,0.00013779472,0.534581,0.00013289148,0.000012197174,0.000105397565,0.0006877981,0.00231154,0.0021932044],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99667144,0.00054364296,0.00020041743,0.0005300401,0.0019116316,0.00014280986],"domain_scores_gemma":[0.9848477,0.0065749055,0.0025396233,0.003326472,0.0025149023,0.00019632073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015558151,0.00096498197,0.0007930517,0.0017057777,0.0005060982,0.0014551412,0.0014454632,0.0011898278,0.0018368582],"category_scores_gemma":[0.019004283,0.0005744119,0.0011982522,0.00073722686,0.00088493933,0.0028530592,0.0016235141,0.0015334557,0.00065894204],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011831963,0.00030926793,0.047797345,0.0011608697,0.00027523155,0.0028073485,0.003304529,0.2212868,0.21213022,0.04676464,0.0048531783,0.4581274],"study_design_scores_gemma":[0.000026468875,0.00022738992,0.0024409017,0.00008007587,0.00006969078,0.0007588575,0.00023091133,0.8093362,0.16322325,0.013102764,0.010445088,0.000058364818],"about_ca_topic_score_codex":0.0028547866,"about_ca_topic_score_gemma":0.0053279684,"teacher_disagreement_score":0.0028547866,"about_ca_system_score_codex":0.0013662565,"about_ca_system_score_gemma":0.001061891,"threshold_uncertainty_score":0.009912968},"labels":[],"label_agreement":null},{"id":"W3147133761","doi":"10.1109/icse.2012.6227207","title":"Recovering traceability links between an API and its learning resources","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":100,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Computer science; Traceability; Documentation; Ambiguity; Code (set theory); Source code; Context (archaeology); World Wide Web; Software engineering; Programming language","score_opus":0.036843110419604065,"score_gpt":0.2864408961637425,"score_spread":0.24959778574413846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3147133761","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15782416,0.000955445,0.7991022,0.0008900424,0.00014550363,0.0005646501,0.0034157177,0.030170634,0.0069316076],"genre_scores_gemma":[0.4142701,0.00074444636,0.5632396,0.00022921848,0.000085841755,0.0004790584,0.0119887395,0.002765023,0.006198058],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9938513,0.0012924712,0.0007415674,0.0012021259,0.0025967765,0.00031577324],"domain_scores_gemma":[0.9371048,0.028516214,0.0071769976,0.013970737,0.012348001,0.0008832325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047862506,0.0015888051,0.00086656294,0.015543021,0.0019521846,0.003990807,0.0023432192,0.0023011873,0.0023619249],"category_scores_gemma":[0.0653812,0.0009623221,0.00131499,0.0065991026,0.0009160313,0.00876063,0.0052586095,0.002796062,0.0017779979],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030014798,0.0006368796,0.051222302,0.0010129286,0.00019327499,0.0011434206,0.0031702823,0.00882625,0.022255927,0.011310844,0.011101327,0.88882643],"study_design_scores_gemma":[0.00013820105,0.00063062186,0.07920146,0.0009012434,0.00070980797,0.002945132,0.0044273823,0.4756284,0.22333626,0.084052935,0.12758355,0.0004449912],"about_ca_topic_score_codex":0.016007809,"about_ca_topic_score_gemma":0.01563981,"teacher_disagreement_score":0.016007809,"about_ca_system_score_codex":0.0014899501,"about_ca_system_score_gemma":0.004316667,"threshold_uncertainty_score":0.031829298},"labels":[],"label_agreement":null},{"id":"W3147362533","doi":"10.1007/s10664-021-09944-w","title":"Revisiting the VCCFinder approach for the identification of vulnerability-contributing commits","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Fonds National de la Recherche Luxembourg; European Commission","keywords":"Commit; Computer science; Identification (biology); Vulnerability (computing); Replicate; Artificial intelligence; Machine learning; Replication (statistics); Software; Set (abstract data type); Software deployment; Data science; Software engineering; Computer security; Programming language; Database","score_opus":0.03679435711364231,"score_gpt":0.3062734892830032,"score_spread":0.26947913216936087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3147362533","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16784161,0.00087287545,0.8099133,0.0017162459,0.00027042456,0.00056362874,0.0022960694,0.01237735,0.0041485312],"genre_scores_gemma":[0.56664884,0.00014117014,0.4250371,0.00037548452,0.00014973637,0.00031258998,0.003972861,0.00040731585,0.00295492],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9895902,0.003704948,0.00072527974,0.0025834679,0.0028616078,0.00053454813],"domain_scores_gemma":[0.91817474,0.045517813,0.007017333,0.014107426,0.013736703,0.001445949],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011131925,0.0009017745,0.0011822212,0.0079924315,0.0013631061,0.002496773,0.0032541999,0.002480243,0.002007028],"category_scores_gemma":[0.05187663,0.00047837864,0.0009182373,0.0025850858,0.0019034467,0.0038057303,0.0033074599,0.0032979583,0.0013079809],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006035688,0.0011630676,0.112592556,0.00079081394,0.00027608764,0.0007455966,0.0018534225,0.08489836,0.013624346,0.015403167,0.026185578,0.7418635],"study_design_scores_gemma":[0.000036361995,0.00019764443,0.010038457,0.00016976542,0.00004410754,0.0005205762,0.00052113185,0.95103824,0.0116633745,0.018160561,0.0075403643,0.0000694746],"about_ca_topic_score_codex":0.004763141,"about_ca_topic_score_gemma":0.007726825,"teacher_disagreement_score":0.011131925,"about_ca_system_score_codex":0.0012320312,"about_ca_system_score_gemma":0.0031308061,"threshold_uncertainty_score":0.058871984},"labels":[],"label_agreement":null},{"id":"W3147763035","doi":"10.48550/arxiv.2104.00725","title":"Assessing the Exposure of Software Changes: The DiPiDi Approach","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software; Computer science; Environmental science; Programming language","score_opus":0.10854892505905768,"score_gpt":0.22202727033135242,"score_spread":0.11347834527229474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3147763035","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6815523,0.00059967954,0.2909437,0.00077986496,0.00008277148,0.0031006793,0.0024540513,0.0055471347,0.014939822],"genre_scores_gemma":[0.73418576,0.0002112029,0.25997078,0.00022583775,0.000025536101,0.0012766345,0.0013433025,0.00024611692,0.002514858],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97374094,0.007973829,0.0017480018,0.004366391,0.011442017,0.00072871585],"domain_scores_gemma":[0.82720435,0.115348674,0.018298658,0.02035899,0.016344527,0.0024449276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01883557,0.0014959986,0.0008914868,0.007412214,0.0008471992,0.0024338458,0.0025622805,0.001635511,0.0025801396],"category_scores_gemma":[0.11218501,0.00074321154,0.0007886566,0.0028597564,0.0011689083,0.0033960426,0.0047720484,0.0024926374,0.0007889261],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031420346,0.0044665853,0.38008413,0.0019449871,0.0009340564,0.00059584784,0.008270588,0.019124841,0.049125466,0.0055303527,0.004906635,0.5218745],"study_design_scores_gemma":[0.0007244447,0.00890246,0.6045916,0.0004457194,0.001181231,0.0022973171,0.0059068305,0.24373299,0.0936666,0.0128934905,0.025072137,0.00058516115],"about_ca_topic_score_codex":0.0027999815,"about_ca_topic_score_gemma":0.0032458662,"teacher_disagreement_score":0.01883557,"about_ca_system_score_codex":0.0016977473,"about_ca_system_score_gemma":0.0018105686,"threshold_uncertainty_score":0.09961325},"labels":[],"label_agreement":null},{"id":"W3148187518","doi":"","title":"Bridging Program Comprehension Tools by Design Navigation","year":2000,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Program comprehension; Computer science; Suite; Bridging (networking); Software engineering; Human–computer interaction; Visualization; Comprehension; Software; Perspective (graphical); World Wide Web; Software system; Programming language; Artificial intelligence","score_opus":0.033685985196059975,"score_gpt":0.28903196195665043,"score_spread":0.25534597676059045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3148187518","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012784903,0.00074604334,0.9689594,0.0009773023,0.000064217864,0.00018342743,0.00007731569,0.00852752,0.007679905],"genre_scores_gemma":[0.05842012,0.00067409506,0.934282,0.00041025138,0.00004808482,0.00032048946,0.00025818744,0.0021827708,0.0034040255],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9885118,0.006352301,0.0010778673,0.001446926,0.0020906364,0.00052048516],"domain_scores_gemma":[0.9376155,0.043800395,0.0029315546,0.012083735,0.0028110466,0.0007577029],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013164378,0.0019369296,0.0010869644,0.0051229703,0.0011698256,0.0071994467,0.0039409306,0.0045437645,0.009180438],"category_scores_gemma":[0.049198944,0.0021518725,0.0011185138,0.0027728465,0.0040564793,0.013457179,0.011290603,0.0038675168,0.003955672],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002399535,0.00073307415,0.003375647,0.0014444079,0.00005218411,0.0010729835,0.028245818,0.0043133325,0.02430675,0.11153276,0.008604543,0.8160786],"study_design_scores_gemma":[0.000359682,0.0010313059,0.0043403744,0.0028479612,0.00023027394,0.006405543,0.0077779265,0.097794995,0.058811698,0.25543514,0.56438875,0.0005763174],"about_ca_topic_score_codex":0.0010496925,"about_ca_topic_score_gemma":0.001146047,"teacher_disagreement_score":0.013164378,"about_ca_system_score_codex":0.0009579718,"about_ca_system_score_gemma":0.0025000146,"threshold_uncertainty_score":0.06962067},"labels":[],"label_agreement":null},{"id":"W3148779029","doi":"10.1109/tse.2020.2988396","title":"A3: Assisting Android API Migrations Using Code Examples","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Android (operating system); Computer science; Documentation; Application programming interface; Source code; World Wide Web; Java; Operating system","score_opus":0.051130970689089575,"score_gpt":0.26172621043122185,"score_spread":0.2105952397421323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3148779029","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5546994,0.0012483133,0.2816092,0.0018355617,0.00038488398,0.0014112804,0.00393692,0.14308988,0.011784452],"genre_scores_gemma":[0.4817454,0.00029964492,0.5034999,0.00040914243,0.000037165068,0.00050250086,0.005353562,0.0011457752,0.007006929],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988985,0.0002829865,0.00008216212,0.00036856902,0.00030746558,0.000060343365],"domain_scores_gemma":[0.99433446,0.0027524303,0.00055652176,0.00084791373,0.0012266,0.0002820887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012113035,0.0015930943,0.0005193779,0.0013590202,0.00048592917,0.00091277284,0.0020955626,0.0014810023,0.0028282069],"category_scores_gemma":[0.014230021,0.00052816793,0.00062914996,0.00059339503,0.0003572801,0.0017841535,0.0011028511,0.0012501094,0.0018468774],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087762723,0.0020416027,0.037965544,0.0010266372,0.00018982602,0.0008906446,0.0017642888,0.04702195,0.02868453,0.0016409995,0.04380671,0.83408964],"study_design_scores_gemma":[0.00015429554,0.00046286837,0.0076549305,0.00013179424,0.00008792342,0.00041525785,0.00042541773,0.9450511,0.02426322,0.0019700648,0.019317074,0.000066072775],"about_ca_topic_score_codex":0.007638099,"about_ca_topic_score_gemma":0.014184008,"teacher_disagreement_score":0.007638099,"about_ca_system_score_codex":0.00046760633,"about_ca_system_score_gemma":0.0013463187,"threshold_uncertainty_score":0.0151872635},"labels":[],"label_agreement":null},{"id":"W3148957464","doi":"10.1109/icse.2002.1007986","title":"Concern graphs: finding and describing concerns using structural program dependencies","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Source code; Scalability; Java; Usability; Abstraction; Representation (politics); Software engineering; Graph; Software; Programming language; Code (set theory); Software maintenance; Software system; Feature (linguistics); Theoretical computer science; Human–computer interaction; Database","score_opus":0.1272718006804312,"score_gpt":0.3360948408474846,"score_spread":0.2088230401670534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3148957464","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009602152,0.0001752574,0.98145837,0.00028435426,0.0000207884,0.00022602652,0.0011646431,0.0055390787,0.0015293754],"genre_scores_gemma":[0.06694886,0.00051649986,0.92508197,0.0001095602,0.000023620076,0.00034999932,0.0036230006,0.0012773332,0.0020691643],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988682,0.00039263885,0.00013429385,0.00019087881,0.0003386922,0.000075310985],"domain_scores_gemma":[0.99448436,0.003199456,0.00073488464,0.0009349155,0.00052961847,0.0001168247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018826687,0.0013927819,0.0004393005,0.005361699,0.0011325976,0.0022377998,0.0015532242,0.0013321248,0.00297999],"category_scores_gemma":[0.008833308,0.00096600846,0.0013173354,0.0030929646,0.0011246637,0.006389632,0.002003134,0.0016273042,0.00071825495],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029255592,0.00034519014,0.018061182,0.0017038739,0.00019112449,0.0019445735,0.013301255,0.07658074,0.022945361,0.26474157,0.032990273,0.5669023],"study_design_scores_gemma":[0.00008202238,0.00017583824,0.005787326,0.00062724244,0.00028282977,0.0018107295,0.0025449754,0.457037,0.032773085,0.2441078,0.25457692,0.00019419243],"about_ca_topic_score_codex":0.010805057,"about_ca_topic_score_gemma":0.017384207,"teacher_disagreement_score":0.010805057,"about_ca_system_score_codex":0.0007960268,"about_ca_system_score_gemma":0.0019506429,"threshold_uncertainty_score":0.021484375},"labels":[],"label_agreement":null},{"id":"W3149382123","doi":"10.1109/iwsc.2012.6227861","title":"An accurate estimation of the Levenshtein distance using metric trees and Manhattan distance","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Levenshtein distance; Metric (unit); Edit distance; Precision and recall; Computer science; Software; Euclidean distance; Distance measurement; Data mining; Artificial intelligence; Engineering","score_opus":0.036356367830956964,"score_gpt":0.3136686935886217,"score_spread":0.27731232575766473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3149382123","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015640292,0.0007700782,0.98173606,0.000049800088,0.00005159775,0.00002846028,0.000078675825,0.0008287407,0.0008163503],"genre_scores_gemma":[0.23055872,0.00082693313,0.7654505,0.000048306192,0.000101309226,0.00008343314,0.0005584826,0.00027012778,0.002102252],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957497,0.00077253085,0.00024292771,0.0006699452,0.0023685717,0.00019620349],"domain_scores_gemma":[0.9925283,0.0031461874,0.0009039403,0.0009704526,0.0022827385,0.0001684317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018699354,0.001083182,0.0011091183,0.0062544807,0.00073474296,0.0016775349,0.001451437,0.0010212556,0.001106194],"category_scores_gemma":[0.014864605,0.00052406883,0.00067563454,0.0039092214,0.00079583947,0.0041651325,0.0011907507,0.0011216663,0.0013123057],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026465874,0.00012530407,0.0150226075,0.0004432631,0.00028465016,0.0003369405,0.00065311824,0.0808067,0.051339943,0.03706311,0.004498345,0.8091614],"study_design_scores_gemma":[0.000020644287,0.00035811705,0.012774081,0.00008457815,0.000083394516,0.0017083813,0.0002861542,0.8725506,0.053572547,0.03603717,0.022309296,0.00021495495],"about_ca_topic_score_codex":0.0035461613,"about_ca_topic_score_gemma":0.003467203,"teacher_disagreement_score":0.0062544807,"about_ca_system_score_codex":0.0009369146,"about_ca_system_score_gemma":0.0007898784,"threshold_uncertainty_score":0.009889245},"labels":[],"label_agreement":null},{"id":"W3150619302","doi":"10.1109/msr.2012.6224307","title":"Mining challenge 2012: The Android platform","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Android (operating system); Computer science; Data science; Android application; Software; World Wide Web; Android app; Operating system","score_opus":0.0470146614936186,"score_gpt":0.27769995712329315,"score_spread":0.23068529562967455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3150619302","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27799794,0.023302488,0.058924284,0.13969584,0.019815393,0.0039642113,0.3101973,0.05857144,0.10753121],"genre_scores_gemma":[0.31565055,0.0054292576,0.07335573,0.010313168,0.0047546425,0.0029186432,0.5044967,0.0070421514,0.07603917],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98139006,0.0035669194,0.0012708688,0.0023825339,0.010121907,0.0012676932],"domain_scores_gemma":[0.9628354,0.012011923,0.002086949,0.007971374,0.0118770925,0.0032173235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01229016,0.0016870679,0.0019121899,0.004468327,0.0031267179,0.0054685404,0.0030521774,0.0044922023,0.0060127084],"category_scores_gemma":[0.045250468,0.00093327067,0.0016011724,0.0040991427,0.0012826609,0.0060056387,0.005940205,0.0044573126,0.007826001],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006115736,0.00039756246,0.007931789,0.00089916977,0.0001226234,0.00056871754,0.00072601467,0.0012992934,0.002432014,0.004931754,0.8961368,0.08394278],"study_design_scores_gemma":[0.00037624393,0.00053391524,0.036889207,0.0004125557,0.00010788434,0.0012090917,0.0014631175,0.0222713,0.008482962,0.008283962,0.919772,0.00019764759],"about_ca_topic_score_codex":0.018764345,"about_ca_topic_score_gemma":0.027241152,"teacher_disagreement_score":0.018764345,"about_ca_system_score_codex":0.0019020276,"about_ca_system_score_gemma":0.005241481,"threshold_uncertainty_score":0.064997375},"labels":[],"label_agreement":null},{"id":"W3150633448","doi":"10.1109/user.2012.6226578","title":"Revisiting bug triage and resolution practices","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Triage; Computer science; Process (computing); Software bug; Categorization; Computer security; Data science; Software; Artificial intelligence; Operating system","score_opus":0.055471890546457715,"score_gpt":0.33548610964797404,"score_spread":0.2800142191015163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3150633448","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9084408,0.0069329618,0.04396455,0.020267975,0.0005457325,0.0014334114,0.0006377251,0.0018836383,0.015893238],"genre_scores_gemma":[0.9475317,0.0038016685,0.04163594,0.0008897658,0.00010676359,0.0008175948,0.00060664513,0.0004736682,0.004136235],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9418345,0.025348645,0.007849315,0.0068519874,0.015642373,0.0024732507],"domain_scores_gemma":[0.6036084,0.22464749,0.05253824,0.029564887,0.084422745,0.005218198],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06740471,0.000850451,0.00074123987,0.015362636,0.0041360594,0.0078411335,0.003278134,0.001847813,0.0029532292],"category_scores_gemma":[0.3185821,0.001482169,0.00089590024,0.009064644,0.003919289,0.012092317,0.004669952,0.0035226871,0.0007541882],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015818469,0.0004440721,0.14405544,0.0022313234,0.00011098086,0.0008752629,0.3876092,0.0009760799,0.0073821372,0.0050017512,0.010944677,0.44021085],"study_design_scores_gemma":[0.00014093898,0.0012503191,0.43285334,0.006567586,0.00027950958,0.0021168902,0.38565272,0.00838434,0.008286472,0.0074099903,0.14663771,0.00042012773],"about_ca_topic_score_codex":0.028213678,"about_ca_topic_score_gemma":0.033503983,"teacher_disagreement_score":0.06740471,"about_ca_system_score_codex":0.011938881,"about_ca_system_score_gemma":0.01372192,"threshold_uncertainty_score":0.35647446},"labels":[],"label_agreement":null},{"id":"W3150814957","doi":"","title":"Lags in the release, adoption, and propagation of npm vulnerability fixes","year":2022,"lang":"en","type":"article","venue":"Institutional Repositories DataBase (IRDB)","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Vulnerability (computing); Software package; Business; Vulnerability assessment; Computer science; Software; Computer security; Operating system; Psychology","score_opus":0.017861238144364185,"score_gpt":0.2611345482604494,"score_spread":0.24327331011608522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3150814957","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9934348,0.00033923643,0.001143308,0.0008041063,0.000040415598,0.000017861312,0.00078304956,0.000078966536,0.0033583941],"genre_scores_gemma":[0.9986519,0.00008939904,0.00021445897,0.00004764421,0.000011534817,0.0000070795595,0.00030111437,0.000009863956,0.00066704146],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9961926,0.0005729537,0.00044697025,0.0008522425,0.001120708,0.00081451004],"domain_scores_gemma":[0.8316406,0.09810429,0.04830928,0.006540386,0.010135664,0.005269779],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072582546,0.00019483003,0.00024244656,0.0020382274,0.00048848207,0.0031229197,0.000785505,0.0012219731,0.006565717],"category_scores_gemma":[0.09880078,0.00045754397,0.0002658285,0.0018068512,0.0009931703,0.003911649,0.0021536949,0.0029498672,0.0009590865],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008863172,0.0003079371,0.9277138,0.00018528198,0.00016358613,0.0003322473,0.0034724362,0.006182794,0.0031336402,0.009756449,0.003318119,0.044547338],"study_design_scores_gemma":[0.000034156117,0.0004055774,0.9791724,0.00011477438,0.00006670303,0.00019494198,0.0044829133,0.0069823596,0.001600982,0.003145253,0.0037413873,0.00005853175],"about_ca_topic_score_codex":0.013370652,"about_ca_topic_score_gemma":0.01730951,"teacher_disagreement_score":0.013370652,"about_ca_system_score_codex":0.0011879959,"about_ca_system_score_gemma":0.0016482295,"threshold_uncertainty_score":0.03838575},"labels":[],"label_agreement":null},{"id":"W3151388396","doi":"10.1109/msr.2012.6224276","title":"Inferring semantically related words from software context","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; WordNet; Java; Natural language processing; Code (set theory); Program comprehension; Software; Context (archaeology); Information retrieval; Software maintenance; Artificial intelligence; Precision and recall; Programming language; Software development; Software system","score_opus":0.018509735782409758,"score_gpt":0.25706644804473966,"score_spread":0.2385567122623299,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3151388396","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5922318,0.0053473883,0.37184703,0.0007051657,0.00018767099,0.00050959265,0.0068678334,0.010792214,0.011511391],"genre_scores_gemma":[0.7424081,0.0016566182,0.24488854,0.00017533163,0.0000804586,0.00016522895,0.008523552,0.0004470185,0.0016552452],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982273,0.00041289668,0.0002177544,0.00056519505,0.00045555457,0.00012133897],"domain_scores_gemma":[0.99515593,0.0027406523,0.0005725153,0.0005834103,0.0008489598,0.000098574405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009118401,0.0014760857,0.0009179522,0.01001271,0.0010720019,0.0015893369,0.0007881273,0.0013143752,0.0028547663],"category_scores_gemma":[0.011158326,0.0005549809,0.0012519021,0.0039101015,0.00061315123,0.0051205605,0.0026297036,0.000946441,0.002151629],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087519136,0.00035249515,0.0742622,0.0025076126,0.00029643546,0.0016747359,0.0045993393,0.009722772,0.126572,0.014437815,0.010887857,0.7538117],"study_design_scores_gemma":[0.00022272064,0.00084469584,0.11087937,0.00092719414,0.0013153778,0.010370501,0.010631803,0.49011967,0.16816449,0.101855315,0.1042843,0.00038453878],"about_ca_topic_score_codex":0.0047093714,"about_ca_topic_score_gemma":0.0073166424,"teacher_disagreement_score":0.01001271,"about_ca_system_score_codex":0.0006378605,"about_ca_system_score_gemma":0.0014155161,"threshold_uncertainty_score":0.009550154},"labels":[],"label_agreement":null},{"id":"W3151410988","doi":"10.1109/icse.2012.6227187","title":"Asking and answering questions about unfamiliar APIs: An exploratory study","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Question answering; Ask price; Process (computing); World Wide Web; Application programming interface; Human–computer interaction; Data science; Artificial intelligence; Programming language","score_opus":0.03202337111645748,"score_gpt":0.3016258249175677,"score_spread":0.2696024538011102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3151410988","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9961073,0.000070698385,0.0024908225,0.00018551998,0.000007690006,0.00041661205,0.00004592419,0.000027647708,0.0006477118],"genre_scores_gemma":[0.98812175,0.00029289836,0.008296271,0.0006966271,0.0000304374,0.0014440521,0.000111808135,0.000056792313,0.00094936],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9860984,0.009379537,0.00075217447,0.0011553812,0.0013420787,0.0012723728],"domain_scores_gemma":[0.88614625,0.09719143,0.004926479,0.0023772523,0.0054551205,0.0039035005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020386001,0.0012341852,0.001176586,0.0017419683,0.004072779,0.0033442178,0.0023371363,0.0031456205,0.0015405407],"category_scores_gemma":[0.07792933,0.0015041678,0.0006101049,0.0010788265,0.004046803,0.0041791354,0.004055474,0.0033376382,0.0005449044],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023913049,0.0020066223,0.03006908,0.00048135722,0.000032396514,0.001934584,0.94610757,0.0001796424,0.008191743,0.0002802491,0.00051294896,0.009964622],"study_design_scores_gemma":[0.00021946411,0.0049126814,0.05226617,0.00041646493,0.00006945969,0.002186891,0.91894454,0.0019969165,0.0050191684,0.0009546779,0.012830354,0.00018323539],"about_ca_topic_score_codex":0.0015541007,"about_ca_topic_score_gemma":0.002639439,"teacher_disagreement_score":0.020386001,"about_ca_system_score_codex":0.001101556,"about_ca_system_score_gemma":0.001809994,"threshold_uncertainty_score":0.10781276},"labels":[],"label_agreement":null},{"id":"W3151667270","doi":"10.1109/asonam49781.2020.9381381","title":"Identifying Social Networks of Programmers using Text Mining for Code Similarity Detection","year":2020,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Code (set theory); Source code; Similarity (geometry); Plagiarism detection; Data mining; Information retrieval; Programming language; Artificial intelligence; Image (mathematics)","score_opus":0.11690088930779048,"score_gpt":0.33976406107768775,"score_spread":0.22286317176989728,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3151667270","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8276722,0.0009702603,0.14812669,0.0005738918,0.000119231954,0.0009286728,0.013300448,0.0020686276,0.0062401146],"genre_scores_gemma":[0.85011977,0.0004388702,0.13037252,0.000072571514,0.00013101251,0.00075663644,0.015091525,0.00006909894,0.0029478983],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99861515,0.0002910766,0.00015167333,0.00045163164,0.0003885443,0.00010191123],"domain_scores_gemma":[0.9963909,0.0017944331,0.00077939813,0.00033503058,0.00050074584,0.00019947578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008563114,0.00077716116,0.00048621857,0.010562368,0.00070346845,0.00093041494,0.00073492015,0.0007391784,0.0010018388],"category_scores_gemma":[0.004380092,0.00019485288,0.00067688903,0.004579597,0.00029566942,0.0015353953,0.0009433905,0.0004618871,0.00076416164],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073213346,0.0012571502,0.26552618,0.0011995318,0.000402746,0.0015192366,0.002167203,0.013643722,0.041073076,0.0051074363,0.0155657455,0.6518059],"study_design_scores_gemma":[0.00008426191,0.0005840737,0.26988128,0.0001589529,0.0002373558,0.0023726076,0.0035159816,0.65297204,0.02847192,0.014873294,0.02672669,0.000121562436],"about_ca_topic_score_codex":0.0025920318,"about_ca_topic_score_gemma":0.0049425424,"teacher_disagreement_score":0.010562368,"about_ca_system_score_codex":0.00042237298,"about_ca_system_score_gemma":0.00055470376,"threshold_uncertainty_score":0.0051538944},"labels":[],"label_agreement":null},{"id":"W3151855275","doi":"10.1109/iwsc.2012.6227876","title":"Towards qualitative comparison of Simulink model clone detection approaches","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; General Motors of Canada","keywords":"Adaptability; clone (Java method); Computer science; Relevance (law); Visualization; Plan (archaeology); Software engineering; Machine learning; Data mining; Artificial intelligence","score_opus":0.2096295790966588,"score_gpt":0.39586194101826205,"score_spread":0.18623236192160325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3151855275","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058252353,0.00024861583,0.93453234,0.00017415009,0.00003988478,0.0002682081,0.00022567796,0.0023776123,0.0038811308],"genre_scores_gemma":[0.51595736,0.00034006932,0.48059374,0.00007198859,0.000014659442,0.00064084213,0.0005603586,0.0005242516,0.001296707],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9907348,0.004083559,0.00062097027,0.0005463772,0.0037814425,0.00023277236],"domain_scores_gemma":[0.94480354,0.034091905,0.0038938816,0.0057210266,0.011097296,0.00039241536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012377207,0.0010429228,0.00060274836,0.0037897415,0.0005284668,0.0027782037,0.0021448387,0.0008931967,0.0045859814],"category_scores_gemma":[0.054009534,0.0005206932,0.0006562183,0.0012254559,0.0014817301,0.0032770871,0.0018941838,0.00081898214,0.0005624024],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00209925,0.0010286258,0.018780533,0.0039790785,0.00035132424,0.00051658164,0.0074657043,0.2260292,0.13852936,0.09646784,0.0029790537,0.50177354],"study_design_scores_gemma":[0.00026264368,0.001257128,0.009262164,0.0005795241,0.00019016874,0.00034715448,0.0037278587,0.75098276,0.17363171,0.04424631,0.015315603,0.00019694846],"about_ca_topic_score_codex":0.0017286283,"about_ca_topic_score_gemma":0.0015774462,"teacher_disagreement_score":0.012377207,"about_ca_system_score_codex":0.0019028009,"about_ca_system_score_gemma":0.0013537864,"threshold_uncertainty_score":0.0654577},"labels":[],"label_agreement":null},{"id":"W3152068601","doi":"10.1109/icse.2012.6227138","title":"Temporal analysis of API usage concepts","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Application programming interface; Documentation; Software; Reuse; Software documentation; Set (abstract data type); Software engineering; Software development; Usage data; Programming language; World Wide Web; Database; Software development process","score_opus":0.028364928366075633,"score_gpt":0.32227386587443396,"score_spread":0.29390893750835834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3152068601","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7468885,0.0016610369,0.23769398,0.0003490561,0.00004435688,0.0003466745,0.0033156646,0.0012853616,0.008415391],"genre_scores_gemma":[0.8942419,0.00043475817,0.10126431,0.00005392903,0.0000286333,0.00035651188,0.0021761246,0.00013855491,0.0013052359],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99658585,0.0005881227,0.00033054338,0.00072981225,0.0015446708,0.00022107858],"domain_scores_gemma":[0.98350024,0.008680525,0.0024484848,0.0014892607,0.003536227,0.00034519358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024421108,0.0004480469,0.00044537144,0.00771895,0.00066186173,0.0012399424,0.0009775017,0.00046862796,0.0015077896],"category_scores_gemma":[0.01868556,0.0003346456,0.00081384764,0.00608552,0.0006463043,0.0023055617,0.0010309058,0.0008278283,0.0002802091],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007418191,0.00048404594,0.24193057,0.0013398167,0.00032086083,0.001936216,0.0088306265,0.01599286,0.05212874,0.03954756,0.0040120515,0.6327348],"study_design_scores_gemma":[0.00009098687,0.00073079683,0.30772886,0.00043951304,0.00041023747,0.0064488538,0.005719193,0.5221874,0.042009402,0.06695085,0.04704158,0.00024230331],"about_ca_topic_score_codex":0.0070746234,"about_ca_topic_score_gemma":0.006635217,"teacher_disagreement_score":0.00771895,"about_ca_system_score_codex":0.0008587316,"about_ca_system_score_gemma":0.001135827,"threshold_uncertainty_score":0.014066875},"labels":[],"label_agreement":null},{"id":"W3152301934","doi":"10.1109/tse.2021.3070269","title":"Software Batch Testing to Save Build Test Resources and to Reduce Feedback Time","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Commit; Test case; Test (biology); Reliability engineering; Database; Machine learning","score_opus":0.015497159235501757,"score_gpt":0.23784210009981144,"score_spread":0.22234494086430967,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3152301934","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2635239,0.0023043354,0.6229098,0.0012149945,0.00053524104,0.0014121749,0.001570835,0.08558472,0.02094407],"genre_scores_gemma":[0.66028893,0.00037598272,0.32423618,0.00045386355,0.00009240934,0.000651737,0.0021870537,0.0049784207,0.0067354306],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.991063,0.0022846307,0.0005890718,0.0015702965,0.0037030736,0.0007899482],"domain_scores_gemma":[0.9482468,0.024068533,0.0031782947,0.015726913,0.0070002754,0.001779087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008526418,0.0021493153,0.0011871787,0.0022182495,0.00065137004,0.0022167924,0.004647088,0.000997464,0.009535929],"category_scores_gemma":[0.03448448,0.0012295317,0.001192333,0.0017093591,0.0012839638,0.004589143,0.0019735969,0.0030600517,0.0030828773],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003085216,0.001591555,0.030173708,0.0014886972,0.00045448256,0.00048665857,0.00094231835,0.19172874,0.079805136,0.015041976,0.03126182,0.64393973],"study_design_scores_gemma":[0.0008402917,0.00294765,0.019349461,0.0002996165,0.0003420017,0.00070092786,0.00044538843,0.8174607,0.09863373,0.020700391,0.03797436,0.00030547334],"about_ca_topic_score_codex":0.008485492,"about_ca_topic_score_gemma":0.009800539,"teacher_disagreement_score":0.009535929,"about_ca_system_score_codex":0.0018998402,"about_ca_system_score_gemma":0.0035903822,"threshold_uncertainty_score":0.045092523},"labels":[],"label_agreement":null},{"id":"W3152332785","doi":"10.1016/b978-0-12-411519-4.00006-9","title":"Latent Dirichlet Allocation","year":2015,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":344,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Topic model; Documentation; Data science; Commit; Variety (cybernetics); Subject-matter expert; Software; Interpretability; Set (abstract data type); Source code; Subject (documents); Information retrieval; World Wide Web; Artificial intelligence; Database","score_opus":0.03667132724998003,"score_gpt":0.26734237749471246,"score_spread":0.23067105024473242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3152332785","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00092293724,0.006370271,0.92784464,0.0015764197,0.00070217514,0.00005431096,0.00070711225,0.0022885918,0.05953353],"genre_scores_gemma":[0.053491965,0.014103972,0.60178703,0.0013239288,0.0016181653,0.0005459327,0.007026443,0.0032432561,0.31685928],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998946,0.0004420171,0.00005756203,0.00026236518,0.00024312877,0.00004892149],"domain_scores_gemma":[0.99887854,0.0006726543,0.000029453777,0.00024155746,0.00013155252,0.000046207115],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017022978,0.0013747927,0.0012663516,0.0015462412,0.00080312986,0.0035897852,0.0012493154,0.0014227168,0.042907227],"category_scores_gemma":[0.0049427557,0.0008909281,0.0012178145,0.0026023793,0.00095136155,0.0029386,0.0027136675,0.0030696413,0.03646136],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000036561443,0.00003732453,0.00019694895,0.00026810664,0.000065756394,0.000058993475,0.00017263411,0.015708514,0.0012932258,0.16860592,0.12375495,0.6898011],"study_design_scores_gemma":[0.000016659822,0.000018765077,0.00044020248,0.0002451394,0.000049466318,0.00029611913,0.0000862731,0.11723716,0.002391383,0.5473821,0.33177692,0.00005980975],"about_ca_topic_score_codex":0.0019740749,"about_ca_topic_score_gemma":0.0034342897,"teacher_disagreement_score":0.042907227,"about_ca_system_score_codex":0.0010152381,"about_ca_system_score_gemma":0.0012339014,"threshold_uncertainty_score":0.14353901},"labels":[],"label_agreement":null},{"id":"W3152918650","doi":"10.1007/s10664-023-10314-x","title":"Evaluating pre-trained models for user feedback analysis in software engineering: a study on classification of app-reviews","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Machine learning; Artificial intelligence; F1 score; Macro; Set (abstract data type); Task (project management); Binary classification; Support vector machine; Data mining; Engineering","score_opus":0.15107932163450608,"score_gpt":0.40137155726348567,"score_spread":0.2502922356289796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3152918650","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9693755,0.0017337397,0.024027567,0.00028633358,0.0001826134,0.00026617874,0.00088440033,0.0015513311,0.0016923426],"genre_scores_gemma":[0.9757111,0.00029312042,0.018400788,0.0001234232,0.000066437264,0.00015067695,0.0032201395,0.00013873282,0.0018956979],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9910953,0.004404254,0.00080830796,0.0018732453,0.0014177454,0.00040098626],"domain_scores_gemma":[0.8748576,0.103799924,0.0030447373,0.005029435,0.011851862,0.0014164327],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011124551,0.0015995802,0.0014535597,0.0025069974,0.00084640924,0.0023439897,0.0021041897,0.0023733014,0.0009961292],"category_scores_gemma":[0.054604802,0.0005244246,0.0011359736,0.0014390876,0.0005410224,0.0028307673,0.0011602737,0.0024410402,0.0013794728],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0055603716,0.007036302,0.24858257,0.0014760012,0.0014684917,0.00044318187,0.0030640182,0.085010566,0.011288842,0.00084322965,0.018163668,0.61706275],"study_design_scores_gemma":[0.00013180343,0.0018336348,0.058121856,0.00017957963,0.0004365423,0.0003125734,0.00070439413,0.9259763,0.0089081945,0.0008217448,0.0024829125,0.000090574635],"about_ca_topic_score_codex":0.012213194,"about_ca_topic_score_gemma":0.013390184,"teacher_disagreement_score":0.98887545,"about_ca_system_score_codex":0.00222015,"about_ca_system_score_gemma":0.0020111306,"threshold_uncertainty_score":0.058832943},"labels":[],"label_agreement":null},{"id":"W3153043709","doi":"10.5430/air.v10n1p34","title":"A study of quality prediction for large-scale open source software projects","year":2021,"lang":"en","type":"article","venue":"Artificial Intelligence Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Deliverable; Quality (philosophy); Resolution (logic); Software; Scale (ratio); Computer science; Open source software; Open source; Product (mathematics); Data science; Data mining; Process management; Business; Artificial intelligence; Engineering; Systems engineering; Mathematics","score_opus":0.2997543026205094,"score_gpt":0.472607513630679,"score_spread":0.1728532110101696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3153043709","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98989797,0.0001783159,0.008879846,0.00017976588,0.00000924927,0.00003655402,0.00031445926,0.00008390768,0.00041998716],"genre_scores_gemma":[0.99627066,0.00006686271,0.0027931812,0.000009415607,0.000010160301,0.000024548624,0.0006581383,0.0000142953995,0.00015274099],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99495643,0.0016010132,0.00057053554,0.001053556,0.0014358496,0.00038263455],"domain_scores_gemma":[0.9052291,0.06514869,0.01630236,0.0031509548,0.008213415,0.0019555192],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009157206,0.00075160316,0.00061974267,0.0052435575,0.0005905182,0.0015537278,0.0009534237,0.0008889253,0.00064510247],"category_scores_gemma":[0.049877293,0.000392385,0.0008792162,0.005135853,0.0005905024,0.0020433443,0.00085453165,0.0012742899,0.0002435856],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009580292,0.00020122688,0.95569247,0.00005199836,0.00011451061,0.00019249259,0.00039755515,0.020994605,0.0004360904,0.00036560895,0.00048298755,0.020974644],"study_design_scores_gemma":[0.000010077016,0.00012737345,0.5906325,0.000018928507,0.000034900757,0.00013464793,0.00032367674,0.40717024,0.00053953915,0.00064452324,0.00033910418,0.000024493605],"about_ca_topic_score_codex":0.013436792,"about_ca_topic_score_gemma":0.006464494,"teacher_disagreement_score":0.013436792,"about_ca_system_score_codex":0.0012699462,"about_ca_system_score_gemma":0.0007016172,"threshold_uncertainty_score":0.048428535},"labels":[],"label_agreement":null},{"id":"W3155762676","doi":"10.5753/jserd.2021.548","title":"Mining Experts from Source Code Analysis: An Empirical Evaluation","year":2021,"lang":"en","type":"article","venue":"Journal of Software Engineering Research and Development","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Java; Source code; Quality (philosophy); Identification (biology); World Wide Web; Software; Software development; Microservices; Software quality; Code refactoring; Software engineering; Space (punctuation); Data science; Operating system; Cloud computing","score_opus":0.09872674669870646,"score_gpt":0.37527940236157586,"score_spread":0.2765526556628694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3155762676","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92496556,0.006381978,0.056512326,0.00067561097,0.00014471066,0.0015874527,0.0034324036,0.0023086765,0.0039913543],"genre_scores_gemma":[0.9126058,0.0010740132,0.07429034,0.00020155952,0.0001238345,0.0007673595,0.0099531105,0.00018938216,0.00079457846],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95455277,0.025949528,0.0041447706,0.004777524,0.009507877,0.0010675633],"domain_scores_gemma":[0.73813856,0.22127801,0.0087727355,0.011471129,0.017664883,0.002674754],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.042624943,0.0019545755,0.0018141208,0.010013041,0.0010736218,0.0024816333,0.0031338288,0.0028388465,0.0011434567],"category_scores_gemma":[0.14024417,0.0005703274,0.0013701656,0.004945789,0.001275811,0.0040617227,0.0027933072,0.0013221141,0.0007835295],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0057587163,0.005988089,0.40283647,0.004844,0.0017733442,0.00053403096,0.0024530096,0.03570002,0.006008957,0.0015084534,0.015252162,0.5173428],"study_design_scores_gemma":[0.001537199,0.004207343,0.21920973,0.0008479935,0.0014107268,0.0012807372,0.0025993602,0.7428896,0.012348544,0.0030988683,0.010405554,0.00016429402],"about_ca_topic_score_codex":0.0035857586,"about_ca_topic_score_gemma":0.0039029242,"teacher_disagreement_score":0.95737505,"about_ca_system_score_codex":0.0010393962,"about_ca_system_score_gemma":0.002110738,"threshold_uncertainty_score":0.22542489},"labels":[],"label_agreement":null},{"id":"W3156120596","doi":"10.1109/iccq51190.2021.9392939","title":"Striffs: Architectural Component Diagrams for Code Reviews","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Code review; Code (set theory); Scope (computer science); Component (thermodynamics); Software engineering; Relation (database); Static program analysis; Key (lock); Programming language; KPI-driven code analysis; Software development; Software; Database; Computer security; Set (abstract data type)","score_opus":0.050285057948109904,"score_gpt":0.30992359464093017,"score_spread":0.2596385366928203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3156120596","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027726695,0.00094250264,0.9094478,0.0017864418,0.00094475056,0.0012747719,0.006808516,0.054667298,0.02135527],"genre_scores_gemma":[0.026147336,0.0011381817,0.94011736,0.00050809846,0.00026434116,0.0015835796,0.008486385,0.009662515,0.01209229],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9924233,0.0032591857,0.0009893335,0.0006063354,0.0024899477,0.00023185459],"domain_scores_gemma":[0.9616548,0.020748902,0.0028024393,0.003874339,0.009913022,0.0010065519],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008728138,0.0018298161,0.00068284315,0.00834335,0.0011447988,0.0041888305,0.0024094011,0.0019262991,0.038375396],"category_scores_gemma":[0.053890113,0.001245552,0.0014013421,0.003543586,0.0008500966,0.004610909,0.003277198,0.0019576764,0.013279047],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003648194,0.000118305856,0.003108613,0.0029788485,0.00009369417,0.001020573,0.0037347064,0.012340434,0.006627781,0.170464,0.2996655,0.4994827],"study_design_scores_gemma":[0.00007643537,0.000072899595,0.0008134822,0.0006627426,0.000037973772,0.0006525881,0.0003447016,0.025782894,0.003416913,0.060471393,0.9075861,0.00008184543],"about_ca_topic_score_codex":0.004062711,"about_ca_topic_score_gemma":0.0074322554,"teacher_disagreement_score":0.038375396,"about_ca_system_score_codex":0.0012481593,"about_ca_system_score_gemma":0.0047325357,"threshold_uncertainty_score":0.12837851},"labels":[],"label_agreement":null},{"id":"W3157144053","doi":"10.5281/zenodo.4726287","title":"Test Smell Detection Tools: A Systematic Mapping Study","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Test (biology); Computer science; Artificial intelligence; Biology","score_opus":0.04614354301624474,"score_gpt":0.2486654630564941,"score_spread":0.20252192004024935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3157144053","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5216744,0.23060168,0.09249376,0.0102472985,0.0012572628,0.04161889,0.07978168,0.001910342,0.020414654],"genre_scores_gemma":[0.6629694,0.066330105,0.14557579,0.0051035658,0.00032002496,0.05679488,0.05886911,0.0008915125,0.0031456507],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.93945616,0.02379492,0.01490471,0.0052292864,0.015017104,0.0015979001],"domain_scores_gemma":[0.57645756,0.29517123,0.034302123,0.02028489,0.07062076,0.0031634774],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07902111,0.0010046609,0.0021610116,0.048202027,0.002179904,0.0040520434,0.0030082532,0.0018485686,0.0035169304],"category_scores_gemma":[0.31163803,0.00091318,0.0031977042,0.03481728,0.0015353324,0.005436933,0.0063779866,0.0020426016,0.0010954844],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080013304,0.0010470373,0.1325968,0.19668566,0.0033460588,0.0019360367,0.050530326,0.0012887478,0.0028680658,0.009261706,0.042085394,0.55755395],"study_design_scores_gemma":[0.0010460395,0.0020101252,0.20782559,0.32112247,0.011798869,0.0028293629,0.08581763,0.0043939976,0.006786409,0.012267979,0.3434513,0.0006503356],"about_ca_topic_score_codex":0.005690103,"about_ca_topic_score_gemma":0.013493743,"teacher_disagreement_score":0.9209789,"about_ca_system_score_codex":0.004871484,"about_ca_system_score_gemma":0.024796301,"threshold_uncertainty_score":0.41790855},"labels":[],"label_agreement":null},{"id":"W3157562347","doi":"10.1007/s10664-020-09926-4","title":"The nature of build changes","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Dependency (UML); Java; Plug-in; Software evolution; Software maintenance; Software engineering; Software; Legacy system; Software system; Programming language; Software construction","score_opus":0.015523738207438768,"score_gpt":0.28122853624808597,"score_spread":0.2657047980406472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3157562347","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92870617,0.004487102,0.044640798,0.0005996328,0.00041706327,0.00037563432,0.006436094,0.004360026,0.009977406],"genre_scores_gemma":[0.9500392,0.0012953589,0.031373445,0.00027144846,0.00014843738,0.00025838756,0.011236277,0.001002339,0.004375127],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9893287,0.0014855823,0.0010953286,0.0023625987,0.0052690743,0.0004587083],"domain_scores_gemma":[0.91399974,0.044053458,0.013724881,0.010989114,0.01615919,0.0010735363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041509187,0.0006121863,0.0004947178,0.007766322,0.0009189095,0.0024457113,0.0009802697,0.00082215597,0.0011469126],"category_scores_gemma":[0.052033357,0.00060680055,0.0006578013,0.003947315,0.00062824966,0.0030186165,0.0015568562,0.0011588394,0.00089960935],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004380662,0.00012512384,0.52808917,0.0012977221,0.0003291909,0.0015328777,0.0066593457,0.0047991956,0.01937775,0.0021197838,0.012144219,0.4230875],"study_design_scores_gemma":[0.000023938748,0.00025796858,0.86913306,0.00062502356,0.00028214743,0.0029498243,0.0025609327,0.021845143,0.022822864,0.0027936993,0.076543555,0.00016186271],"about_ca_topic_score_codex":0.0031490403,"about_ca_topic_score_gemma":0.0053048353,"teacher_disagreement_score":0.007766322,"about_ca_system_score_codex":0.0008036872,"about_ca_system_score_gemma":0.00063956284,"threshold_uncertainty_score":0.02195239},"labels":[],"label_agreement":null},{"id":"W3158412293","doi":"10.1145/3463274.3463342","title":"DABT: A Dependency-aware Bug Triaging Method","year":2021,"lang":"en","type":"article","venue":"Evaluation and Assessment in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software regression; Dependency (UML); Software bug; Security bug; Software; Blocking (statistics); Process (computing)","score_opus":0.03924068237700611,"score_gpt":0.3814679397277233,"score_spread":0.3422272573507172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158412293","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004825071,0.00018173047,0.9741072,0.00030860078,0.000086446074,0.00028719637,0.0004181748,0.01858227,0.0012033177],"genre_scores_gemma":[0.062381253,0.00017901885,0.9307789,0.00021418744,0.00006646167,0.00040541322,0.0014181694,0.001851789,0.0027047608],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974584,0.0005528953,0.00024837285,0.00058556267,0.0009639033,0.00019097058],"domain_scores_gemma":[0.9954485,0.002181303,0.00064791087,0.00057717744,0.00091921753,0.0002259169],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023228321,0.0018504716,0.0010415327,0.0045388583,0.0010469237,0.0013245852,0.0025306947,0.0011675851,0.005912902],"category_scores_gemma":[0.009691575,0.0010625655,0.002250424,0.0021052964,0.0009153255,0.0025041038,0.002184441,0.0020112032,0.0017767502],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002915968,0.00030882025,0.0071915705,0.0008232474,0.00018291708,0.0005883574,0.00050640217,0.03907168,0.020376727,0.015407263,0.046598606,0.8686529],"study_design_scores_gemma":[0.0003654498,0.00023722331,0.0020428556,0.00012771232,0.00024653116,0.0011232025,0.00022467822,0.89111346,0.020516891,0.03681283,0.047054306,0.00013475325],"about_ca_topic_score_codex":0.0055764713,"about_ca_topic_score_gemma":0.006891945,"teacher_disagreement_score":0.005912902,"about_ca_system_score_codex":0.0011491795,"about_ca_system_score_gemma":0.0043880013,"threshold_uncertainty_score":0.019780636},"labels":[],"label_agreement":null},{"id":"W3159118496","doi":"10.1016/j.infsof.2021.106603","title":"How do developers discuss and support new programming languages in technical Q&amp;A site? An empirical study of Go, Swift, and Rust in Stack Overflow","year":2021,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Swift; Java; Rust (programming language); Second-generation programming language; Resource (disambiguation); Software; Stack (abstract data type); Fourth-generation programming language; Programming language; Software development; World Wide Web; Data science; Software engineering; Programming paradigm; Fifth-generation programming language","score_opus":0.019263389115637304,"score_gpt":0.3074252941101488,"score_spread":0.2881619049945115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159118496","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9951106,0.000049563027,0.00056920946,0.00105386,0.000011369714,0.000041617135,0.000023186454,0.000048726357,0.0030918927],"genre_scores_gemma":[0.9971636,0.00006935723,0.00084863626,0.0004341772,0.000014409984,0.000041116127,0.000049083013,0.000053312346,0.0013263528],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9854181,0.008546796,0.0007489465,0.00083187665,0.003197387,0.0012567589],"domain_scores_gemma":[0.61551183,0.2639586,0.07173433,0.009107702,0.018914778,0.0207727],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.020433798,0.00036377765,0.00041900156,0.0029128375,0.003961969,0.0045627686,0.0018733348,0.003313435,0.004397812],"category_scores_gemma":[0.19441502,0.0010559829,0.00034828187,0.0017862294,0.0042517073,0.012242791,0.0043060537,0.004920591,0.0007899006],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005191681,0.002725697,0.6147394,0.00023249665,0.000079869256,0.0008930847,0.33417058,0.00034372625,0.0022364988,0.0041976627,0.004083667,0.035778146],"study_design_scores_gemma":[0.00018554479,0.0011799564,0.5412061,0.0003217474,0.00012340149,0.00047781126,0.43286586,0.0032477411,0.0013498607,0.0044867597,0.01437607,0.00017920119],"about_ca_topic_score_codex":0.013763639,"about_ca_topic_score_gemma":0.030275712,"teacher_disagreement_score":0.996038,"about_ca_system_score_codex":0.0028786696,"about_ca_system_score_gemma":0.005618944,"threshold_uncertainty_score":0.108065546},"labels":[],"label_agreement":null},{"id":"W3159568147","doi":"10.1145/3439769","title":"Automatic API Usage Scenario Documentation from Technical Q&amp;A Sites","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Saskatchewan; University of Calgary","funders":"","keywords":"Documentation; Computer science; Internal documentation; Software documentation; Code (set theory); Application programming interface; World Wide Web; Coding (social sciences); Software engineering; Information retrieval; Software; Programming language; Software development; Software development process; Set (abstract data type); Software construction","score_opus":0.09554814590246685,"score_gpt":0.34831518452753624,"score_spread":0.2527670386250694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159568147","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4118486,0.005461448,0.4010733,0.0031374185,0.0010252381,0.002712585,0.07070371,0.067154065,0.03688372],"genre_scores_gemma":[0.35724676,0.001569885,0.5228461,0.00042196142,0.00035233473,0.001989808,0.10310376,0.002861694,0.009607626],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9928108,0.002342563,0.00072729844,0.0012295871,0.0026306838,0.00025903148],"domain_scores_gemma":[0.9571361,0.01507274,0.0065373634,0.00342732,0.016623188,0.0012033618],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044480837,0.0012753566,0.0007988072,0.015090626,0.0010801706,0.0026383821,0.0012276828,0.0014711525,0.0036825566],"category_scores_gemma":[0.033807054,0.0008255937,0.000772122,0.006754344,0.00037185717,0.0028198077,0.002278447,0.0013615141,0.0051872195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044517778,0.00046763723,0.045422927,0.0028943159,0.00018965435,0.002868245,0.005199231,0.0051129553,0.021620719,0.003463793,0.17980647,0.7325089],"study_design_scores_gemma":[0.00027101955,0.00061689684,0.12918353,0.002453806,0.00035698497,0.0038791527,0.009761102,0.2983918,0.055255786,0.014106393,0.48511687,0.00060667767],"about_ca_topic_score_codex":0.002592572,"about_ca_topic_score_gemma":0.0057570906,"teacher_disagreement_score":0.015090626,"about_ca_system_score_codex":0.0010168692,"about_ca_system_score_gemma":0.00218717,"threshold_uncertainty_score":0.023523986},"labels":[],"label_agreement":null},{"id":"W3160121929","doi":"10.1109/saner50967.2021.00061","title":"MSR4ML: Reconstructing Artifact Traceability in Machine Learning Repositories","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Traceability; Computer science; Artifact (error); Software engineering; Software versioning; Source code; Commit; Software evolution; Software development; Requirements traceability; Software; Process (computing); Database; Software construction; Programming language; Artificial intelligence","score_opus":0.017582682260829176,"score_gpt":0.25896396172902053,"score_spread":0.24138127946819135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3160121929","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068036735,0.00061346556,0.87499046,0.0005746336,0.00006469654,0.00027438859,0.0036928693,0.050505728,0.0012469889],"genre_scores_gemma":[0.25781348,0.00036124137,0.7229639,0.00011233346,0.00004365149,0.00030671782,0.013786214,0.0026528335,0.0019595881],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9938461,0.0015257292,0.00072310556,0.0013918107,0.0021901529,0.00032314248],"domain_scores_gemma":[0.9664835,0.01241207,0.0057154545,0.0115811275,0.003232179,0.0005756901],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.007004681,0.0012230212,0.00082619657,0.009680835,0.0009649437,0.0036154191,0.003239738,0.0016341374,0.0015044529],"category_scores_gemma":[0.039891038,0.000988605,0.0015163891,0.00539687,0.0011039332,0.0053803734,0.004460903,0.0020103296,0.0011979715],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003694671,0.000553561,0.08742762,0.0011469701,0.00037123947,0.0012603287,0.0030882074,0.07459222,0.01651766,0.01707511,0.013978514,0.7836191],"study_design_scores_gemma":[0.00005633442,0.00027524823,0.015443686,0.00028623964,0.00013028466,0.00085496856,0.00077156705,0.87937224,0.043658987,0.033101596,0.025951717,0.00009713768],"about_ca_topic_score_codex":0.005979814,"about_ca_topic_score_gemma":0.008277765,"teacher_disagreement_score":0.9929953,"about_ca_system_score_codex":0.0010335905,"about_ca_system_score_gemma":0.002517016,"threshold_uncertainty_score":0.037044704},"labels":[],"label_agreement":null},{"id":"W3160351645","doi":"10.1109/icse-companion52605.2021.00129","title":"Interactive Graph Exploration for Comprehension of Static Analysis Results","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Codebase; Computer science; Program comprehension; Visualization; Static analysis; Data visualization; Graph; Graphical user interface; Static program analysis; Power graph analysis; Cognitive load; Human–computer interaction; Comprehension; Interactive visualization; Source code; Call graph; Graph drawing; Programming language; Theoretical computer science; Cognition; Software; Data mining; Software development; Software system","score_opus":0.04481574509969042,"score_gpt":0.3214316351998798,"score_spread":0.27661589010018933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3160351645","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038837343,0.0007646276,0.8287944,0.0014356044,0.00020054537,0.000491346,0.0034126753,0.11083552,0.015227931],"genre_scores_gemma":[0.27251562,0.0009256063,0.7019483,0.00045233354,0.00015127879,0.00066856074,0.0045752474,0.013406052,0.005357011],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975508,0.0012365958,0.00017867818,0.00038172377,0.0004977513,0.00015436372],"domain_scores_gemma":[0.9449178,0.048636634,0.0009210282,0.0028621515,0.0018844,0.00077788683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004057302,0.0034742847,0.0013813989,0.004754331,0.0007537172,0.0048726005,0.0019552482,0.002046723,0.04993219],"category_scores_gemma":[0.03606605,0.00076530105,0.0012907038,0.0022165794,0.00096500764,0.0073857387,0.0043547754,0.0016971518,0.0065382216],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00524134,0.0009863393,0.0068100044,0.006382373,0.00034216212,0.002743704,0.028308917,0.019803854,0.100101985,0.050143547,0.100540206,0.6785956],"study_design_scores_gemma":[0.0011241689,0.0010794903,0.009091532,0.0021845559,0.00041251635,0.0021149013,0.0065797702,0.42315707,0.066993326,0.2415616,0.24508739,0.00061373087],"about_ca_topic_score_codex":0.0016810489,"about_ca_topic_score_gemma":0.0023753864,"teacher_disagreement_score":0.04993219,"about_ca_system_score_codex":0.00080379023,"about_ca_system_score_gemma":0.0010515072,"threshold_uncertainty_score":0.16703981},"labels":[],"label_agreement":null},{"id":"W3160394822","doi":"10.1109/saner50967.2021.00045","title":"EnHMM: On the Use of Ensemble HMMs and Stack Traces to Predict the Reassignment of Bug Report Fields","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Concordia University","funders":"","keywords":"Computer science; Eclipse; Hidden Markov model; Stack (abstract data type); Support vector machine; Measure (data warehouse); Field (mathematics); Precision and recall; Artificial intelligence; Data mining; Machine learning; Function (biology); Mathematics; Programming language","score_opus":0.061574191970904475,"score_gpt":0.27540748669718995,"score_spread":0.21383329472628548,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3160394822","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21280798,0.004483253,0.7236095,0.0012587926,0.00062729686,0.0003507009,0.006365128,0.047775965,0.0027214526],"genre_scores_gemma":[0.73180395,0.0011041949,0.24284458,0.0008299892,0.00026563377,0.00024129394,0.016055562,0.0009011051,0.005953652],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985153,0.00037629966,0.00012733963,0.00051644875,0.00029594335,0.00016860318],"domain_scores_gemma":[0.9944074,0.0029123346,0.00046707556,0.00095326855,0.0010006065,0.00025930905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026102506,0.0022825415,0.0014329426,0.0031875777,0.00059688545,0.00091853744,0.0022725272,0.0014166132,0.0010380677],"category_scores_gemma":[0.009495171,0.0007832402,0.0013858237,0.001609207,0.0004419075,0.0020175376,0.0016067429,0.00259577,0.0012684828],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076197623,0.00064106425,0.076953426,0.00032820422,0.0006136357,0.00037926386,0.00041603763,0.26365235,0.006859507,0.0016299518,0.023329008,0.6244355],"study_design_scores_gemma":[0.000025469317,0.000112861344,0.0041214344,0.000030477519,0.000078034645,0.00007723873,0.000040011124,0.9895298,0.002354477,0.0017499849,0.0018507821,0.000029399807],"about_ca_topic_score_codex":0.030503022,"about_ca_topic_score_gemma":0.04466866,"teacher_disagreement_score":0.030503022,"about_ca_system_score_codex":0.0009340126,"about_ca_system_score_gemma":0.0019657773,"threshold_uncertainty_score":0.060650945},"labels":[],"label_agreement":null},{"id":"W3161311896","doi":"10.1145/1454497.1454485","title":"Dynamic analysis of Ada programs for comprehension and quality measurement","year":2008,"lang":"en","type":"article","venue":"ACM SIGAda Ada Letters","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Program comprehension; Schema (genetic algorithms); Software engineering; Software quality; Comprehension; Visualization; Programming language; Call graph; Software; Set (abstract data type); Software system; Data mining; Software development; Information retrieval","score_opus":0.08656111550795259,"score_gpt":0.30134099927986446,"score_spread":0.21477988377191187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3161311896","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4862401,0.0005200614,0.48422778,0.00031639502,0.000040464776,0.00027775503,0.0014675959,0.022960266,0.003949616],"genre_scores_gemma":[0.8229867,0.00022182145,0.17263398,0.000057220343,0.00002892185,0.00023265339,0.0015465791,0.0013536842,0.00093839836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971214,0.0006121988,0.0002170583,0.00044796933,0.0014724619,0.00012884857],"domain_scores_gemma":[0.98769134,0.005245186,0.0016145973,0.0020125133,0.0031887745,0.00024758597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020079662,0.0007988045,0.0006140504,0.0038233686,0.00037046944,0.0014799619,0.00066275493,0.0004273637,0.0010623646],"category_scores_gemma":[0.011613202,0.00032420582,0.0004986863,0.0021891096,0.00050845544,0.0016546864,0.0007021791,0.0009956672,0.00037963747],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007928906,0.0007681139,0.09256607,0.0008536768,0.00019322363,0.0005228175,0.0044496604,0.03668693,0.30511045,0.010313111,0.004247909,0.54349524],"study_design_scores_gemma":[0.00007182438,0.0013365534,0.12699081,0.00016960628,0.00021459533,0.0009471068,0.0012065124,0.57922053,0.25613624,0.012594124,0.020922013,0.00019008541],"about_ca_topic_score_codex":0.0016345389,"about_ca_topic_score_gemma":0.0014748422,"teacher_disagreement_score":0.0038233686,"about_ca_system_score_codex":0.0006305195,"about_ca_system_score_gemma":0.0007212046,"threshold_uncertainty_score":0.010619223},"labels":[],"label_agreement":null},{"id":"W3161504508","doi":"10.1109/icse43902.2021.00135","title":"CodeShovel: Constructing Method-Level Source Code Histories","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Source code; Correctness; Software engineering; Code review; Field (mathematics); Programming language; Software; Code (set theory); Empirical research; Static program analysis; Software development","score_opus":0.05602422699197473,"score_gpt":0.31455817382952383,"score_spread":0.2585339468375491,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3161504508","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14620164,0.0017709192,0.6917556,0.0006877058,0.00018354315,0.0013625494,0.030940393,0.11934503,0.0077526188],"genre_scores_gemma":[0.29872552,0.001081803,0.6338144,0.00017645437,0.00006583735,0.0010653156,0.048652988,0.010769737,0.0056480397],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9969523,0.000619807,0.00033707867,0.0007458461,0.0011899744,0.0001549318],"domain_scores_gemma":[0.9715816,0.013505639,0.003895593,0.005689007,0.0046187607,0.0007095133],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004640053,0.0011586617,0.0005383652,0.009494727,0.0010603857,0.0025452226,0.001529077,0.0009401616,0.0031377382],"category_scores_gemma":[0.03713779,0.0013932621,0.00093591044,0.003792378,0.00094177126,0.005542779,0.0028866273,0.0014851509,0.0022512947],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005930605,0.0002827375,0.13513215,0.0020319696,0.00021898461,0.00094015355,0.010675879,0.013649181,0.013732156,0.019214228,0.045089934,0.75843954],"study_design_scores_gemma":[0.00023434887,0.0006820225,0.088440605,0.0018152124,0.00035224116,0.0020912075,0.0051403777,0.36941838,0.10192806,0.061247755,0.3681449,0.00050494535],"about_ca_topic_score_codex":0.0069063627,"about_ca_topic_score_gemma":0.01274313,"teacher_disagreement_score":0.009494727,"about_ca_system_score_codex":0.0010051075,"about_ca_system_score_gemma":0.0045975633,"threshold_uncertainty_score":0.024539292},"labels":[],"label_agreement":null},{"id":"W3161551626","doi":"10.1109/saner50967.2021.00030","title":"The Usability (or Not) of Refactoring Tools","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Code refactoring; Computer science; Software engineering; Usability; USable; Software; Human–computer interaction; Programming language; World Wide Web","score_opus":0.06920001959988052,"score_gpt":0.31622509905957213,"score_spread":0.2470250794596916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3161551626","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.943181,0.002744161,0.031750526,0.002604937,0.00015811583,0.00037857014,0.000100025194,0.00038542616,0.018697327],"genre_scores_gemma":[0.98791784,0.00053775526,0.010208292,0.0002332502,0.000037544607,0.00011326013,0.00007297255,0.00010055453,0.000778623],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.90669703,0.055680342,0.009140723,0.0024637682,0.023393253,0.0026249245],"domain_scores_gemma":[0.5881299,0.32291177,0.019879663,0.01953898,0.046878994,0.0026608068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0575681,0.0006156367,0.00059629383,0.0029598412,0.0015033932,0.004572699,0.0012713925,0.002069876,0.00075412745],"category_scores_gemma":[0.3096,0.0007330418,0.0009865481,0.0015207719,0.0030529492,0.0065369112,0.0017623496,0.0013498258,0.00031944652],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087930966,0.0005065925,0.3185283,0.0040231077,0.0006142128,0.0016275571,0.23192738,0.001773898,0.019167515,0.0078088273,0.00315215,0.4099911],"study_design_scores_gemma":[0.00029969463,0.0040809033,0.6059203,0.007366109,0.0010907495,0.0056506125,0.24569021,0.015235788,0.018041473,0.015463096,0.08029761,0.0008634617],"about_ca_topic_score_codex":0.0028125434,"about_ca_topic_score_gemma":0.0035695117,"teacher_disagreement_score":0.0575681,"about_ca_system_score_codex":0.0017709208,"about_ca_system_score_gemma":0.002204084,"threshold_uncertainty_score":0.3044529},"labels":[],"label_agreement":null},{"id":"W3161588130","doi":"10.1109/icse-companion52605.2021.00100","title":"CodeShovel: A Reusable and Available Tool for Extracting Source Code Histories","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Source code; Java; Software engineering; Component (thermodynamics); World Wide Web; Code (set theory); Software; Programming language; Database","score_opus":0.03151616611182169,"score_gpt":0.2691346081720582,"score_spread":0.23761844206023652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3161588130","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0063270256,0.00075539853,0.6377249,0.0005383187,0.00024007568,0.00075975363,0.030322904,0.31117156,0.012160123],"genre_scores_gemma":[0.04440165,0.0016832373,0.7469428,0.00038700367,0.0001515965,0.0014095098,0.11006027,0.07667499,0.01828891],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966304,0.00038228685,0.00039383819,0.0005612033,0.0018744047,0.0001578554],"domain_scores_gemma":[0.9781528,0.009178716,0.0027375016,0.004526468,0.0047403057,0.00066418975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031689012,0.0018968849,0.0007987436,0.011712025,0.0016778383,0.0036769034,0.0017808523,0.0015063062,0.019330237],"category_scores_gemma":[0.035927687,0.0019125812,0.0014145748,0.0060367375,0.0010742889,0.005691339,0.0052072853,0.0031381173,0.020101635],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003850042,0.00013982145,0.011354398,0.002771322,0.00016485949,0.0007561438,0.00276957,0.0044319653,0.012586876,0.024479149,0.26160532,0.67855555],"study_design_scores_gemma":[0.00012010824,0.0001336075,0.009016418,0.0015938054,0.000109051216,0.0015500952,0.00079676294,0.050943706,0.05043308,0.03870767,0.8462463,0.00034942443],"about_ca_topic_score_codex":0.0051789214,"about_ca_topic_score_gemma":0.008445735,"teacher_disagreement_score":0.019330237,"about_ca_system_score_codex":0.0010766274,"about_ca_system_score_gemma":0.0069387867,"threshold_uncertainty_score":0.06466609},"labels":[],"label_agreement":null},{"id":"W3162129107","doi":"10.1109/icse-companion52605.2021.00058","title":"A Better Approach to Track the Evolution of Static Code Warnings","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Static analysis; Computer science; Workflow; Code (set theory); Static program analysis; Matching (statistics); Tracking (education); Source code; Detector; Software; Set (abstract data type); Software engineering; Software evolution; Track (disk drive); Plug-in; Programming language; Software system; Software development; Operating system; Database; Software construction","score_opus":0.01933404994448727,"score_gpt":0.2535206605101129,"score_spread":0.23418661056562565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3162129107","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16794705,0.0025135907,0.7857603,0.0014077012,0.00033947185,0.00033524918,0.0028611366,0.03563064,0.0032048712],"genre_scores_gemma":[0.4823422,0.0005643013,0.5084032,0.00039781127,0.00008769234,0.00016676525,0.0045090322,0.0009476501,0.0025813435],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99396825,0.0009900114,0.0005915615,0.002109562,0.0020446635,0.00029602498],"domain_scores_gemma":[0.98327255,0.00431036,0.0034028057,0.0037921797,0.004834688,0.0003874929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042129545,0.0014574461,0.0011209681,0.00897681,0.00074854644,0.0020690705,0.0014648366,0.0019126977,0.001479506],"category_scores_gemma":[0.022925638,0.00066755776,0.0010998694,0.004076984,0.00047391868,0.0029295594,0.001644361,0.0016496292,0.0012355726],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029395267,0.00044487027,0.09926734,0.00085287576,0.00028239,0.00041826422,0.0009883272,0.02089912,0.051990706,0.0039138934,0.012107766,0.80854046],"study_design_scores_gemma":[0.00015015563,0.00084186427,0.10812851,0.00037796,0.0006079784,0.0017267582,0.00084845454,0.719455,0.090412356,0.0142529225,0.06283531,0.00036275218],"about_ca_topic_score_codex":0.007732161,"about_ca_topic_score_gemma":0.010447649,"teacher_disagreement_score":0.00897681,"about_ca_system_score_codex":0.00068163994,"about_ca_system_score_gemma":0.0022385302,"threshold_uncertainty_score":0.022280514},"labels":[],"label_agreement":null},{"id":"W3162755873","doi":"10.1016/j.infsof.2021.106636","title":"Requirements engineering: Foundation for software quality (REFSQ2020)","year":2021,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Foundation (evidence); Software quality; Software engineering; Quality (philosophy); Engineering; Computer science; Systems engineering; Software; Software development; Programming language; Geography","score_opus":0.0263946985204725,"score_gpt":0.30104182476449726,"score_spread":0.27464712624402476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3162755873","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003516502,0.014143075,0.62154216,0.043013472,0.008956166,0.0007530452,0.011631657,0.016100613,0.28034332],"genre_scores_gemma":[0.0497884,0.017895203,0.7206187,0.008214632,0.004901508,0.0016445043,0.0532933,0.0113393385,0.13230439],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9837818,0.0051106033,0.0011028323,0.00093347323,0.007993269,0.0010780278],"domain_scores_gemma":[0.96352535,0.008456927,0.0020965415,0.0060760113,0.018524956,0.0013202467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020883271,0.0017562254,0.0013647452,0.0035680349,0.0012862043,0.0069185887,0.0020551893,0.0050678207,0.04450421],"category_scores_gemma":[0.06009734,0.0011578057,0.0014956446,0.0036550574,0.0011879026,0.006679898,0.004434302,0.004426947,0.050637733],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008280748,0.00011375119,0.0004805016,0.00076112757,0.00004014073,0.000053825916,0.00012088284,0.0029686103,0.0019388977,0.25880745,0.423049,0.31158304],"study_design_scores_gemma":[0.00007390021,0.00009731469,0.0008727002,0.0016711025,0.000043769065,0.0001468325,0.000058245412,0.009644686,0.002255966,0.18625188,0.7988326,0.000051001676],"about_ca_topic_score_codex":0.0066582565,"about_ca_topic_score_gemma":0.0028038267,"teacher_disagreement_score":0.04450421,"about_ca_system_score_codex":0.0029899008,"about_ca_system_score_gemma":0.007952178,"threshold_uncertainty_score":0.14888138},"labels":[],"label_agreement":null},{"id":"W3163202066","doi":"10.1109/tse.2021.3082068","title":"An Empirical Study of Type-Related Defects in Python Projects","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Python (programming language); Computer science; Programming language; Artificial intelligence","score_opus":0.025406470675922346,"score_gpt":0.2945407086543776,"score_spread":0.26913423797845526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3163202066","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.994288,0.00023661212,0.0018440222,0.00040152043,0.000023698847,0.00019949398,0.0004966014,0.00009081378,0.0024192785],"genre_scores_gemma":[0.99571496,0.00015844184,0.001933995,0.00014662085,0.000019312201,0.0003808391,0.00067821564,0.0000660596,0.00090146635],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.956293,0.016468026,0.0051749256,0.0049427026,0.01534587,0.0017755631],"domain_scores_gemma":[0.43530193,0.31167158,0.15736502,0.019454086,0.06721427,0.008993046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034676667,0.00047196238,0.00051609246,0.0057401387,0.0017033606,0.0029589788,0.0018746298,0.0018389908,0.002532884],"category_scores_gemma":[0.318782,0.00066147454,0.00052884495,0.005409111,0.002998511,0.0059979553,0.0033795324,0.0030026897,0.00090512395],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002706387,0.00061424705,0.9561905,0.00028846512,0.0000630402,0.00039988465,0.013401199,0.00034617263,0.0003946144,0.00086054014,0.0023308946,0.024839826],"study_design_scores_gemma":[0.000060637103,0.0010190569,0.9527838,0.00041370618,0.00008172021,0.0010758527,0.028973626,0.0039531062,0.0012977354,0.0012946129,0.008951855,0.00009424749],"about_ca_topic_score_codex":0.0038519094,"about_ca_topic_score_gemma":0.0042889123,"teacher_disagreement_score":0.034676667,"about_ca_system_score_codex":0.0025440012,"about_ca_system_score_gemma":0.0029715213,"threshold_uncertainty_score":0.18338996},"labels":[],"label_agreement":null},{"id":"W3163466556","doi":"10.1109/icse-companion52605.2021.00056","title":"Explainable Just-In-Time Bug Prediction: Are We There Yet?","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Interpretability; Computer science; Predictive modelling; Software bug; Machine learning; Task (project management); Data mining; Artificial intelligence; Software; Programming language; Engineering","score_opus":0.0249446827098048,"score_gpt":0.25515023628695566,"score_spread":0.23020555357715086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3163466556","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19854799,0.0046407366,0.7533332,0.022858469,0.00047680867,0.00029894203,0.00223616,0.014173236,0.0034344485],"genre_scores_gemma":[0.77311856,0.0014417241,0.21716163,0.0022454208,0.0002568486,0.00018377196,0.003290703,0.0011633335,0.001137913],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.991015,0.003603472,0.00068077166,0.0018881805,0.002307135,0.0005053397],"domain_scores_gemma":[0.8498189,0.09465277,0.014449729,0.028950408,0.0108672585,0.001260998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013269067,0.0018596126,0.0011770244,0.0025154229,0.00084690866,0.003721385,0.0028814427,0.0025896586,0.0027591544],"category_scores_gemma":[0.10876437,0.0009958185,0.0016588105,0.0017847831,0.0019292424,0.01495257,0.0030715389,0.005014536,0.0007998076],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000795262,0.00092658255,0.26176447,0.0022780672,0.0010465991,0.0012857241,0.0066686785,0.054140072,0.009162165,0.03652652,0.017104384,0.6083015],"study_design_scores_gemma":[0.00026732887,0.0008474711,0.08707376,0.0015050872,0.0012206994,0.0018536192,0.0036761526,0.66816044,0.018927269,0.17486544,0.04101206,0.0005906532],"about_ca_topic_score_codex":0.0082311025,"about_ca_topic_score_gemma":0.009629923,"teacher_disagreement_score":0.013269067,"about_ca_system_score_codex":0.0019756933,"about_ca_system_score_gemma":0.0025566856,"threshold_uncertainty_score":0.070174396},"labels":[],"label_agreement":null},{"id":"W3163555559","doi":"10.1109/icpc52881.2021.00011","title":"Assessing Semantic Frames to Support Program Comprehension Activities","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Program comprehension; Correctness; Natural language processing; Parsing; Software engineering; Artificial intelligence; Programming language; Software system; Software","score_opus":0.04289885128387787,"score_gpt":0.35115703209666954,"score_spread":0.30825818081279166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3163555559","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37963554,0.0004410601,0.59902835,0.0006793776,0.00008417358,0.0011124299,0.0019179326,0.011831465,0.005269659],"genre_scores_gemma":[0.5838296,0.00018003445,0.40978986,0.00008749161,0.000026226488,0.00040935207,0.0042256,0.0006863475,0.0007654649],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9888193,0.006403003,0.0011796395,0.0012622674,0.0018141617,0.000521665],"domain_scores_gemma":[0.907138,0.06592652,0.005670036,0.0067030615,0.013755863,0.00080651895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014587081,0.0019160131,0.0010094672,0.010702474,0.001325915,0.0036086163,0.0013488078,0.002370388,0.0028429276],"category_scores_gemma":[0.102118686,0.0006753181,0.0013421106,0.0040277066,0.0012207237,0.008837601,0.0025078002,0.0015710391,0.0011367422],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015641061,0.0014299965,0.10098914,0.0020142659,0.00044542603,0.0007425244,0.015679177,0.04181643,0.06862308,0.02642811,0.010904309,0.7293634],"study_design_scores_gemma":[0.0002654645,0.0011790611,0.044195615,0.0005043005,0.00068917195,0.00055061054,0.009517531,0.7475329,0.13508405,0.03986737,0.020319104,0.00029479002],"about_ca_topic_score_codex":0.009392382,"about_ca_topic_score_gemma":0.009504293,"teacher_disagreement_score":0.014587081,"about_ca_system_score_codex":0.0015723835,"about_ca_system_score_gemma":0.003647118,"threshold_uncertainty_score":0.07714474},"labels":[],"label_agreement":null},{"id":"W3163683940","doi":"10.1109/saner50967.2021.00073","title":"A Testing Approach While Re-engineering Legacy Systems: An Industrial Case Study","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Process (computing); Computer science; Automation; Quality (philosophy); Reliability engineering; Legacy system; Redevelopment; Risk analysis (engineering); Reliability (semiconductor); System testing; Systems engineering; Engineering; Software engineering; Software; Business","score_opus":0.14075779375369252,"score_gpt":0.2958936475241719,"score_spread":0.1551358537704794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3163683940","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8984645,0.0007215031,0.078558944,0.0010949107,0.000049473467,0.0005762496,0.00017580444,0.0004721672,0.019886462],"genre_scores_gemma":[0.93335736,0.0003603823,0.0623338,0.00013684324,0.000019003957,0.000118841934,0.00017583912,0.000063092826,0.003434869],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970009,0.0016382167,0.00014889563,0.00025869577,0.00067956344,0.00027365363],"domain_scores_gemma":[0.9905433,0.006546267,0.00057294953,0.0009840778,0.0009700221,0.00038337117],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030581972,0.0006696331,0.00030180294,0.0012416212,0.0013953062,0.0014302867,0.0017281552,0.00223962,0.0013165415],"category_scores_gemma":[0.0069196606,0.00026900068,0.000566997,0.0012769306,0.0011992845,0.0013564439,0.0009555266,0.0011825052,0.00030337923],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013137064,0.0093712695,0.090953834,0.0015024628,0.00017382085,0.044144873,0.016316056,0.21000385,0.046192575,0.05427323,0.008407449,0.5173469],"study_design_scores_gemma":[0.00071591476,0.007376806,0.04568986,0.0008009825,0.00037403248,0.017143024,0.015064362,0.72334224,0.07460775,0.024730586,0.08988845,0.0002659199],"about_ca_topic_score_codex":0.008141688,"about_ca_topic_score_gemma":0.01075284,"teacher_disagreement_score":0.008141688,"about_ca_system_score_codex":0.0013817232,"about_ca_system_score_gemma":0.0011659925,"threshold_uncertainty_score":0.016188622},"labels":[],"label_agreement":null},{"id":"W3163691159","doi":"10.1007/s10664-021-09954-8","title":"On using Stack Overflow comment-edit pairs to recommend code maintenance changes","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Commit; Code (set theory); Source code; Code review; Code smell; Data mining; Programming language; Point (geometry); Interface (matter); Static program analysis; Set (abstract data type); Software; Database; Software development; Software quality; Operating system","score_opus":0.05172537609133042,"score_gpt":0.3103960695023846,"score_spread":0.25867069341105414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3163691159","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6899592,0.0023995775,0.2691909,0.0035326476,0.0007645626,0.0009781377,0.0037086813,0.0146079855,0.014858222],"genre_scores_gemma":[0.82551813,0.00039953622,0.16050836,0.0008274178,0.00025990387,0.00017710842,0.0046341782,0.00031234868,0.007362953],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99739516,0.0008049625,0.00020820869,0.0005193134,0.00087826484,0.00019404183],"domain_scores_gemma":[0.9648567,0.026067542,0.0009411076,0.0018713893,0.005596375,0.00066681154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042079166,0.0013161892,0.0013134367,0.004612102,0.0014779993,0.0022057425,0.0019275398,0.0032296313,0.0057501565],"category_scores_gemma":[0.035258695,0.0005875561,0.0006272203,0.0020291284,0.00050177414,0.0041373144,0.0014917126,0.0018164009,0.0025999425],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038211448,0.0016281037,0.06542669,0.00035813978,0.00036105196,0.0002430553,0.0007034273,0.023196135,0.013257322,0.0023673028,0.049761556,0.8388761],"study_design_scores_gemma":[0.00042605025,0.00079648755,0.02809656,0.00010883169,0.00024355961,0.00026712572,0.0009078412,0.9459252,0.010749635,0.005375653,0.0069768997,0.00012633442],"about_ca_topic_score_codex":0.03999728,"about_ca_topic_score_gemma":0.069588736,"teacher_disagreement_score":0.03999728,"about_ca_system_score_codex":0.0009600267,"about_ca_system_score_gemma":0.0019618296,"threshold_uncertainty_score":0.07952893},"labels":[],"label_agreement":null},{"id":"W3170475911","doi":"10.22215/etd/2020-13954","title":"An Empirical Study Investigating the Predictors of Software Metric Correlation in Application Code and Test Code.","year":2020,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Code (set theory); Metric (unit); Computer science; Software metric; Software; Code coverage; Dead code; Regression testing; Programming language; Redundant code; Software quality; Software construction; Software development; Engineering; Code generation; Operating system; Set (abstract data type)","score_opus":0.025783966780925573,"score_gpt":0.3243294239617642,"score_spread":0.2985454571808386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3170475911","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9960277,0.00028112,0.00083596347,0.00026393624,0.00001812023,0.000043077493,0.0007435936,0.000025719606,0.001760781],"genre_scores_gemma":[0.9972778,0.00012560292,0.00059372257,0.000052408996,0.00001627809,0.00007429827,0.0011612135,0.000021283138,0.00067746],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9903066,0.0051063634,0.00065173063,0.0010399957,0.002403259,0.000492067],"domain_scores_gemma":[0.54825544,0.36695218,0.046225015,0.009864664,0.020537829,0.008164794],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013114098,0.00034461526,0.00032450794,0.0022588926,0.0005490635,0.0014250482,0.0009583981,0.0008258723,0.004361125],"category_scores_gemma":[0.1883071,0.0003350363,0.00044430414,0.0038190265,0.0011155237,0.0020919612,0.0010617818,0.0022839697,0.0015217403],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020874635,0.00030922994,0.9908632,0.000034586577,0.000083957544,0.00004629811,0.0006111651,0.00017890472,0.00007240782,0.00028150523,0.0009822837,0.0063277343],"study_design_scores_gemma":[0.000020842515,0.00044828694,0.9936367,0.000036414283,0.00004630425,0.00017690055,0.0014554968,0.0025466124,0.00015209606,0.00031932935,0.0011466694,0.000014235622],"about_ca_topic_score_codex":0.0021942337,"about_ca_topic_score_gemma":0.0022772115,"teacher_disagreement_score":0.013114098,"about_ca_system_score_codex":0.0007871407,"about_ca_system_score_gemma":0.00089437526,"threshold_uncertainty_score":0.06935477},"labels":[],"label_agreement":null},{"id":"W3173462463","doi":"10.1109/icpc52881.2021.00051","title":"Warning-Introducing Commits vs Bug-Introducing Commits: A tool, statistical models, and a preliminary user study","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Commit; Computer science; Logistic regression; Odds; Statistical model; Software bug; Predictive power; Software; Software engineering; Database; Machine learning; Programming language","score_opus":0.021134144775830963,"score_gpt":0.2742574760266623,"score_spread":0.2531233312508313,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173462463","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95621395,0.0002845114,0.034833252,0.00038974397,0.000048879996,0.00048709658,0.0022175282,0.0043961196,0.0011288903],"genre_scores_gemma":[0.94195825,0.0001205662,0.05185221,0.0001812433,0.00003727846,0.0006709825,0.00308392,0.0010079016,0.0010876589],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98192847,0.011689023,0.0014128316,0.0020987946,0.0024034681,0.00046738435],"domain_scores_gemma":[0.6283747,0.33034307,0.005694772,0.020742543,0.012921666,0.0019233077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027917728,0.0017783932,0.0013047644,0.0040557343,0.00084732834,0.0029257608,0.0022675248,0.001802539,0.0030893777],"category_scores_gemma":[0.13849919,0.0012515134,0.0013752478,0.0028388114,0.0011150732,0.0047636875,0.00272255,0.003089705,0.0011804884],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0071184356,0.009302169,0.6266943,0.0020783485,0.0010522127,0.0019848817,0.019438982,0.06603808,0.015557938,0.0052103708,0.026601035,0.21892323],"study_design_scores_gemma":[0.00059832976,0.007752472,0.12871182,0.00035560998,0.00032720185,0.0017079765,0.00426531,0.8228398,0.019765722,0.0046067545,0.00857754,0.00049144827],"about_ca_topic_score_codex":0.0043455726,"about_ca_topic_score_gemma":0.0064274203,"teacher_disagreement_score":0.027917728,"about_ca_system_score_codex":0.0009600085,"about_ca_system_score_gemma":0.00068968645,"threshold_uncertainty_score":0.14764482},"labels":[],"label_agreement":null},{"id":"W3173682469","doi":"10.1109/msr52588.2021.00040","title":"Studying the Change Histories of Stack Overflow and GitHub Snippets","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Stack (abstract data type); World Wide Web; Code (set theory); Source code; Information retrieval; Key (lock); Programming language; Computer security","score_opus":0.061319704235336094,"score_gpt":0.27252581574699897,"score_spread":0.21120611151166288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173682469","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9613026,0.0017204197,0.00679625,0.0004709122,0.00014106788,0.00017419325,0.02381851,0.0025055245,0.0030705577],"genre_scores_gemma":[0.8543653,0.0012711176,0.028498892,0.00027054118,0.00015514909,0.000591848,0.10741561,0.0016185171,0.005813063],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9956334,0.0006820657,0.00042697135,0.0011432427,0.0017928518,0.00032157585],"domain_scores_gemma":[0.9758301,0.01175453,0.004293142,0.0025819563,0.004819407,0.0007209647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025794795,0.00092323555,0.00058122765,0.011058718,0.0009105201,0.0015344482,0.0010551392,0.0009863878,0.0009362431],"category_scores_gemma":[0.02891625,0.00043120855,0.00073757686,0.009408678,0.0009502971,0.0027803988,0.0019270766,0.0009883925,0.0009338253],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048666052,0.00026570391,0.7195533,0.0019884077,0.00046584176,0.003098147,0.0071769217,0.00662407,0.01364524,0.0016876587,0.036272816,0.20873518],"study_design_scores_gemma":[0.000046011508,0.00025285344,0.878412,0.00043354908,0.00023939091,0.0027714837,0.0045285556,0.03905965,0.01247915,0.0024346055,0.059175514,0.0001671943],"about_ca_topic_score_codex":0.009202959,"about_ca_topic_score_gemma":0.021817075,"teacher_disagreement_score":0.011058718,"about_ca_system_score_codex":0.00069412915,"about_ca_system_score_gemma":0.0009491647,"threshold_uncertainty_score":0.018298805},"labels":[],"label_agreement":null},{"id":"W3173862657","doi":"","title":"Towards auto-completion on software requirements statements.","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Requirements elicitation; Computer science; Software development; Software requirements specification; Software engineering; Context (archaeology); Software; Software requirements; Quality (philosophy); Software project management; Requirement prioritization; Requirements analysis; Software development process; Software construction; Programming language","score_opus":0.13147481789348378,"score_gpt":0.2605053479608819,"score_spread":0.1290305300673981,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173862657","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006958748,0.0002458691,0.9892424,0.00031435472,0.0000519127,0.0002537128,0.00016533826,0.0012297407,0.0015379589],"genre_scores_gemma":[0.07682472,0.00042076196,0.9141037,0.0002535259,0.00015058476,0.00039576794,0.0023062178,0.0006565966,0.004888161],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98336375,0.009566242,0.0009291721,0.0017723795,0.0039737923,0.00039461537],"domain_scores_gemma":[0.9144439,0.06150701,0.003544593,0.008126187,0.011707269,0.00067111064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012884539,0.002117202,0.0010898863,0.003504301,0.0008620758,0.00211298,0.0015285762,0.001402284,0.007625889],"category_scores_gemma":[0.077105224,0.0011685013,0.0034829492,0.0023278212,0.0018655937,0.0039131395,0.0038921505,0.003960419,0.004010406],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048223164,0.00051043654,0.0026383838,0.002550819,0.00034542114,0.0017287916,0.004873889,0.12057395,0.022615224,0.13446423,0.019457135,0.6897595],"study_design_scores_gemma":[0.0001306148,0.00026020088,0.0011733279,0.00040977274,0.00009576163,0.00067874056,0.0008964713,0.6560239,0.020575596,0.28824377,0.031435817,0.000076052784],"about_ca_topic_score_codex":0.002845129,"about_ca_topic_score_gemma":0.0033107968,"teacher_disagreement_score":0.012884539,"about_ca_system_score_codex":0.0008597162,"about_ca_system_score_gemma":0.0019191235,"threshold_uncertainty_score":0.068140805},"labels":[],"label_agreement":null},{"id":"W3173944821","doi":"10.1145/3463274.3463343","title":"Assessing Developer Expertise from the Statistical Distribution of Programming Syntax Patterns","year":2021,"lang":"en","type":"article","venue":"Evaluation and Assessment in Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Task (project management); Syntax; Context (archaeology); Software engineering; Programming language; Data science; Artificial intelligence; Systems engineering; Engineering","score_opus":0.04264145128656933,"score_gpt":0.357883313491921,"score_spread":0.31524186220535166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173944821","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75550115,0.0006291655,0.23651421,0.0003252994,0.000044971443,0.0001430695,0.0013431667,0.0008562924,0.004642776],"genre_scores_gemma":[0.97964305,0.00017604788,0.017988922,0.00005938741,0.000032439057,0.000100927115,0.0013497088,0.00011356991,0.00053586374],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98867315,0.0040612146,0.0009024495,0.0021792229,0.0038493513,0.00033460633],"domain_scores_gemma":[0.7794584,0.18113112,0.013933065,0.010855631,0.012651639,0.001970176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015065602,0.0004979078,0.0006241573,0.006935353,0.0003688675,0.0015753087,0.0005832196,0.0011709041,0.0018434874],"category_scores_gemma":[0.14890976,0.00031658448,0.00045696818,0.0026548277,0.0010025717,0.0028228867,0.0016518224,0.0010519576,0.0010555524],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070430344,0.00022114662,0.7063872,0.00026678227,0.00035341273,0.00020099618,0.0011000755,0.02212478,0.009368991,0.0024783446,0.002964349,0.25382957],"study_design_scores_gemma":[0.00005580106,0.00069774315,0.6962022,0.00012014835,0.00011697637,0.0011864235,0.0008133033,0.26726508,0.009455976,0.02088236,0.00307769,0.00012625384],"about_ca_topic_score_codex":0.0013164079,"about_ca_topic_score_gemma":0.0017923896,"teacher_disagreement_score":0.015065602,"about_ca_system_score_codex":0.0004493401,"about_ca_system_score_gemma":0.0007459579,"threshold_uncertainty_score":0.079675436},"labels":[],"label_agreement":null},{"id":"W3174073333","doi":"10.1109/msr52588.2021.00065","title":"Mea culpa: How developers fix their own simple bugs differently from other developers","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Commit; Java; Computer science; Software bug; Scope (computer science); Software engineering; Statement (logic); Code (set theory); Simple (philosophy); Code refactoring; Security bug; World Wide Web; Software; Programming language; Database; Computer security; Software security assurance","score_opus":0.028283070769276927,"score_gpt":0.23881036685188176,"score_spread":0.21052729608260484,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174073333","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97565335,0.00191121,0.0062641613,0.00047177952,0.00010608183,0.00008537566,0.010303453,0.0020516478,0.0031529828],"genre_scores_gemma":[0.9515812,0.00048454668,0.011029489,0.00023210324,0.000089987225,0.0001264095,0.032858096,0.0007756445,0.0028225193],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.995631,0.0014359873,0.00037931674,0.0013243491,0.00093397195,0.00029537908],"domain_scores_gemma":[0.9623023,0.018879507,0.00860513,0.0057204543,0.0033811652,0.0011114043],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0041838526,0.00073268486,0.00059955753,0.0040622475,0.00078467454,0.0019201976,0.0009984965,0.00086438784,0.0012619392],"category_scores_gemma":[0.03179848,0.0003299765,0.0006117142,0.0033913867,0.00057044864,0.0026222302,0.0014570728,0.0009292526,0.0011713459],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004034252,0.00016617238,0.911855,0.00041410845,0.00036707084,0.00026321306,0.002253905,0.001718607,0.0025970591,0.0006451953,0.0151554905,0.064160705],"study_design_scores_gemma":[0.00007106673,0.00036444716,0.9291951,0.00014385696,0.00024223982,0.0013976592,0.0028923007,0.024468277,0.0038139957,0.0023417897,0.034951188,0.00011801795],"about_ca_topic_score_codex":0.0066429973,"about_ca_topic_score_gemma":0.014630063,"teacher_disagreement_score":0.9958162,"about_ca_system_score_codex":0.0006687852,"about_ca_system_score_gemma":0.00066985685,"threshold_uncertainty_score":0.022126615},"labels":[],"label_agreement":null},{"id":"W3174274784","doi":"10.48550/arxiv.2105.05981","title":"Assessing Semantic Frames to Support Program Comprehension Activities","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Program comprehension; Correctness; Natural language processing; Software engineering; Parsing; Artificial intelligence; Software development; Programming language; Software; Software system","score_opus":0.10332762653772086,"score_gpt":0.25829985106384973,"score_spread":0.15497222452612885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174274784","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37832403,0.00048354446,0.5999229,0.0007290341,0.00008957563,0.0010243302,0.002037321,0.012057491,0.0053318157],"genre_scores_gemma":[0.58327144,0.0002006515,0.40960282,0.00009366272,0.000029582157,0.0003944664,0.0048198975,0.0007741119,0.0008133631],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9891994,0.0061830054,0.0011064109,0.0012211282,0.001788572,0.0005014017],"domain_scores_gemma":[0.9085146,0.06533511,0.005465094,0.006539904,0.013327796,0.0008174627],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013909043,0.0018370023,0.0009876813,0.010445914,0.0013388289,0.0037708322,0.0013217302,0.002429886,0.0027996246],"category_scores_gemma":[0.10138521,0.00067856966,0.0013004849,0.004139474,0.0012466463,0.008903529,0.0025551477,0.00160543,0.0011888241],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015682231,0.0014556955,0.10096666,0.002034637,0.00043225457,0.0007170082,0.015053587,0.043594465,0.067733355,0.028915187,0.012225775,0.7253031],"study_design_scores_gemma":[0.00026420137,0.0010937885,0.0412735,0.00048018177,0.0006510178,0.0004934444,0.009299361,0.747768,0.13226989,0.045173693,0.020957168,0.0002757215],"about_ca_topic_score_codex":0.009564046,"about_ca_topic_score_gemma":0.009288355,"teacher_disagreement_score":0.013909043,"about_ca_system_score_codex":0.0015789022,"about_ca_system_score_gemma":0.0035918735,"threshold_uncertainty_score":0.07355893},"labels":[],"label_agreement":null},{"id":"W3175255846","doi":"10.1109/mobilesoft52590.2021.00010","title":"An Empirical Study on the Impact of Refactoring on Quality Metrics in Android Applications","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Android (operating system); Software quality; Software engineering; Software maintenance; Empirical research; Software; Software development; Programming language; Operating system","score_opus":0.1436329687129933,"score_gpt":0.46751016015766345,"score_spread":0.32387719144467014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3175255846","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99633694,0.000276014,0.0021701842,0.00018399266,0.0000055886308,0.000047360823,0.00019901371,0.00001643027,0.0007645877],"genre_scores_gemma":[0.99777466,0.000097616525,0.001759977,0.00002670426,0.000005862244,0.00002852291,0.00017755695,0.000007205748,0.000121960846],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9802949,0.011078519,0.0014054071,0.0022079288,0.00427399,0.0007392102],"domain_scores_gemma":[0.36175665,0.5643207,0.044762455,0.011480416,0.01561278,0.002067017],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.022386175,0.00043090383,0.00033522802,0.002663218,0.00050865806,0.0011796689,0.0009852579,0.00092099345,0.001211955],"category_scores_gemma":[0.20078328,0.00034104867,0.0007458438,0.003430085,0.001369737,0.0026167524,0.0010392665,0.001921761,0.00019804669],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024098855,0.0005993902,0.96140075,0.00023508478,0.00021929601,0.00018815752,0.0022534407,0.0026703204,0.00080277666,0.0008604188,0.00025356602,0.030275876],"study_design_scores_gemma":[0.000030039018,0.00087750645,0.97897834,0.000094425886,0.00016220429,0.00020248232,0.001789572,0.014333809,0.0018206135,0.0007276581,0.00095266954,0.0000307373],"about_ca_topic_score_codex":0.0054099886,"about_ca_topic_score_gemma":0.005110979,"teacher_disagreement_score":0.9776138,"about_ca_system_score_codex":0.0014955093,"about_ca_system_score_gemma":0.0014141647,"threshold_uncertainty_score":0.11839086},"labels":[],"label_agreement":null},{"id":"W3175852614","doi":"10.1109/msr52588.2021.00018","title":"An Empirical Study of Developer Discussions on Low-Code Software Development Challenges","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Personalization; Computer science; Event (particle physics); Software; Software development; Empirical research; Code (set theory); World Wide Web; Code review; Software engineering; Data science; Static program analysis; Programming language","score_opus":0.13734326350757375,"score_gpt":0.26043258565926714,"score_spread":0.12308932215169338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3175852614","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9975707,0.00009176799,0.00032723148,0.00025192808,0.000009522902,0.00005577062,0.00018108892,0.000011152326,0.0015008535],"genre_scores_gemma":[0.9965469,0.00026095397,0.00081231433,0.0002409927,0.000041990745,0.0003446308,0.0005630425,0.000028748538,0.0011605539],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9924349,0.0040089884,0.00061457464,0.0007309308,0.0014697836,0.0007407818],"domain_scores_gemma":[0.83049196,0.12484649,0.02502494,0.0024443266,0.012849594,0.0043427204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008686684,0.00041639546,0.00046746613,0.0035863924,0.0018536671,0.0027449788,0.0007526019,0.0013635507,0.0027612073],"category_scores_gemma":[0.07721413,0.00051854464,0.00029288154,0.0030689978,0.0014130345,0.0046797656,0.0022657975,0.001692508,0.0008246774],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037595734,0.00054736645,0.56584835,0.00068217656,0.000039578026,0.0007177137,0.39286533,0.00019350412,0.003284033,0.0012691842,0.0031008513,0.031075917],"study_design_scores_gemma":[0.000043543834,0.00049551966,0.52733195,0.0004181131,0.000029867278,0.0004567071,0.44858834,0.0023038741,0.0012679233,0.00052574783,0.018456424,0.00008201506],"about_ca_topic_score_codex":0.0034298876,"about_ca_topic_score_gemma":0.0037844307,"teacher_disagreement_score":0.008686684,"about_ca_system_score_codex":0.0014457075,"about_ca_system_score_gemma":0.001210503,"threshold_uncertainty_score":0.0459401},"labels":[],"label_agreement":null},{"id":"W3176989815","doi":"10.1007/s10664-021-10066-6","title":"Test case selection and prioritization using machine learning: a systematic literature review","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":158,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Mitacs; Huawei Technologies; Canada Research Chairs","keywords":"Regression testing; Computer science; Machine learning; Prioritization; Artificial intelligence; Feature selection; Process (computing); Software; Test case; Test (biology); Software engineering; Regression analysis; Software development; Engineering; Management science; Software construction","score_opus":0.019469488494802994,"score_gpt":0.28500089700866355,"score_spread":0.26553140851386053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176989815","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010676198,0.9761461,0.0068027824,0.0016495943,0.00025014917,0.0027337333,0.0006866548,0.00006691495,0.0009878216],"genre_scores_gemma":[0.12953013,0.8321085,0.030721243,0.0020184899,0.00027611342,0.0038158,0.0011527241,0.000069865826,0.00030709436],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9552804,0.018287372,0.014796616,0.0027365715,0.008360227,0.00053867436],"domain_scores_gemma":[0.73209107,0.22594377,0.020362787,0.0047323997,0.015695766,0.0011741499],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.05975972,0.0019583036,0.007467157,0.021022012,0.0012616508,0.003955563,0.004611086,0.002447627,0.003002858],"category_scores_gemma":[0.20527458,0.0014213604,0.006923026,0.01496554,0.00170537,0.0054163244,0.0026922813,0.0019000317,0.00039086648],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060780614,0.00033599694,0.006460019,0.53913754,0.007851888,0.0002643355,0.0011868357,0.0010629719,0.00056655996,0.000979983,0.0035188561,0.43802717],"study_design_scores_gemma":[0.00091627456,0.0009945198,0.010880734,0.8951286,0.050811745,0.001148266,0.0023542584,0.0024966463,0.0019649847,0.003350326,0.029751172,0.00020252314],"about_ca_topic_score_codex":0.0056227976,"about_ca_topic_score_gemma":0.015517925,"teacher_disagreement_score":0.94024026,"about_ca_system_score_codex":0.00584488,"about_ca_system_score_gemma":0.027835494,"threshold_uncertainty_score":0.31604338},"labels":[],"label_agreement":null},{"id":"W3177170112","doi":"10.1109/formalise52586.2021.00015","title":"Checking temporal patterns of API usage without code execution","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Programming language; Code (set theory); Coding (social sciences)","score_opus":0.02931014359280479,"score_gpt":0.2879984287300426,"score_spread":0.2586882851372378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3177170112","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36923927,0.0003461477,0.5908143,0.0003222566,0.00017536475,0.00030298144,0.000968524,0.033966247,0.0038649407],"genre_scores_gemma":[0.803439,0.00015218496,0.19006053,0.0002513999,0.000033650536,0.0003446229,0.0012400143,0.0027638264,0.0017148368],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97882926,0.00511698,0.0020856708,0.004223111,0.008271728,0.0014732713],"domain_scores_gemma":[0.9203924,0.035947137,0.011488271,0.020312518,0.010867441,0.0009921853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008188703,0.0013371578,0.0007992896,0.0021258309,0.0007685994,0.0018469917,0.0021767716,0.0011551001,0.0018203763],"category_scores_gemma":[0.0614843,0.0012094255,0.000958343,0.0015266206,0.0018649804,0.0036841037,0.0022669684,0.0019821962,0.00073674094],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020862115,0.0009205046,0.14280096,0.0015467424,0.00053084,0.0018793172,0.0028507665,0.07197247,0.19632436,0.021754127,0.008549392,0.5487843],"study_design_scores_gemma":[0.00026335538,0.0015880431,0.03414621,0.00028919682,0.0003533674,0.0016334074,0.00057026907,0.5886911,0.33293903,0.01941534,0.019849049,0.00026165092],"about_ca_topic_score_codex":0.006456772,"about_ca_topic_score_gemma":0.007257543,"teacher_disagreement_score":0.008188703,"about_ca_system_score_codex":0.00090455357,"about_ca_system_score_gemma":0.003995653,"threshold_uncertainty_score":0.04330653},"labels":[],"label_agreement":null},{"id":"W3177321543","doi":"10.1109/msr52588.2021.00037","title":"On the Use of Dependabot Security Pull Requests","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Merge (version control); Computer science; Secure coding; Computer security; Security bug; Software security assurance; JavaScript; Software; Application security; World Wide Web; Security service; Information security; Information retrieval; Programming language","score_opus":0.055538491442416855,"score_gpt":0.2762570972963118,"score_spread":0.22071860585389497,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3177321543","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9719149,0.0004382453,0.02016762,0.00028025237,0.000026287673,0.00015467357,0.00086133153,0.0033377293,0.0028190108],"genre_scores_gemma":[0.97612303,0.00030255388,0.018140124,0.00010342074,0.000022379916,0.00012619482,0.0022990617,0.00080736494,0.00207597],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.988303,0.0026951774,0.00085310853,0.0018749426,0.0052903495,0.0009834213],"domain_scores_gemma":[0.93187517,0.03752528,0.01475326,0.007318023,0.0065024537,0.0020258268],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068941843,0.0013044216,0.00065429584,0.0048219035,0.001003578,0.0022852567,0.0013337351,0.0012262417,0.0015747487],"category_scores_gemma":[0.05103975,0.0012616412,0.0011182294,0.002325239,0.001251153,0.004358941,0.0029323897,0.0016249408,0.0008669672],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000783506,0.0005459789,0.82446617,0.000772276,0.00038796017,0.003380436,0.007043395,0.033548947,0.019757617,0.003954002,0.004052237,0.10130739],"study_design_scores_gemma":[0.00006673215,0.0010289233,0.60105157,0.00043917057,0.00042803906,0.0044262493,0.0038578277,0.34319124,0.015584548,0.006379604,0.023244929,0.00030117892],"about_ca_topic_score_codex":0.008190009,"about_ca_topic_score_gemma":0.009338,"teacher_disagreement_score":0.008190009,"about_ca_system_score_codex":0.0012960464,"about_ca_system_score_gemma":0.00172189,"threshold_uncertainty_score":0.0364604},"labels":[],"label_agreement":null},{"id":"W3177400593","doi":"10.1109/msr52588.2021.00050","title":"Rollback Edit Inconsistencies in Developer Forum","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Rollback; Computer science; Consistency (knowledge bases); Quality (philosophy); Collaborative editing; World Wide Web; Software quality; Information retrieval; Software; Data science; Database; Programming language; Software development; Artificial intelligence; Database transaction","score_opus":0.017465384677730053,"score_gpt":0.2492278717460679,"score_spread":0.23176248706833785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3177400593","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8813058,0.00073946855,0.10834998,0.0011329502,0.00012581261,0.00069910666,0.0007517875,0.0028968696,0.003998147],"genre_scores_gemma":[0.9258332,0.00018755032,0.07011483,0.00031877198,0.00006184175,0.00037695552,0.0010409337,0.0004533317,0.0016127209],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9174632,0.041492652,0.008415374,0.0063695395,0.024217982,0.0020413299],"domain_scores_gemma":[0.44903904,0.41201556,0.075680524,0.026695665,0.03463901,0.0019301363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06095632,0.0007819612,0.00074587227,0.009505985,0.0023970613,0.0035349121,0.0018860633,0.0017156425,0.0017055381],"category_scores_gemma":[0.31825304,0.0009993674,0.00073824916,0.0048153503,0.0023434032,0.0076011377,0.005347069,0.0015769483,0.0003491965],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077758403,0.00041942773,0.48915642,0.0016464586,0.0002516596,0.0024467697,0.13693757,0.004093565,0.012001339,0.01364834,0.0066165114,0.33200437],"study_design_scores_gemma":[0.00028778947,0.0012064687,0.5637048,0.0034158528,0.00082742865,0.005904882,0.12509094,0.108461425,0.043635435,0.049096867,0.09755672,0.0008114185],"about_ca_topic_score_codex":0.0036587615,"about_ca_topic_score_gemma":0.004363828,"teacher_disagreement_score":0.06095632,"about_ca_system_score_codex":0.0026588014,"about_ca_system_score_gemma":0.0030483173,"threshold_uncertainty_score":0.32237166},"labels":[],"label_agreement":null},{"id":"W3177451191","doi":"10.1109/msr52588.2021.00066","title":"PySStuBs: Characterizing Single-Statement Bugs in Popular Open-Source Python Projects","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Python (programming language); Java; Computer science; Workflow; Open source; Programming language; Software engineering; Software bug; Operating system; World Wide Web; Software; Database","score_opus":0.05650532925885045,"score_gpt":0.2993577384506876,"score_spread":0.24285240919183715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3177451191","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95328254,0.00090500904,0.025916684,0.00046331203,0.00008874644,0.00024683194,0.0049359966,0.011007938,0.0031528675],"genre_scores_gemma":[0.9322769,0.0005752953,0.048455335,0.00030676727,0.000054918106,0.0004668849,0.010821843,0.004032809,0.0030092641],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9933003,0.001072033,0.0009437458,0.0011876863,0.0029263413,0.0005698053],"domain_scores_gemma":[0.9390064,0.025799904,0.020005524,0.004874023,0.00799879,0.0023154558],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004695791,0.0009754391,0.00043629427,0.0063195843,0.0009994031,0.0013267822,0.0012956668,0.00074489805,0.0015998207],"category_scores_gemma":[0.04621242,0.0006471025,0.00063060276,0.0047138412,0.0014521831,0.0034304678,0.003143798,0.0011672076,0.0007285234],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005922121,0.00032768998,0.8103641,0.0016180777,0.00021713243,0.0024746002,0.00892545,0.003474697,0.015099183,0.0031183318,0.023404893,0.13038366],"study_design_scores_gemma":[0.00016191625,0.0008580277,0.83123034,0.00082862575,0.00038132002,0.0070575103,0.0076290625,0.06284622,0.029369963,0.009712933,0.049548123,0.00037590382],"about_ca_topic_score_codex":0.004843249,"about_ca_topic_score_gemma":0.008094029,"teacher_disagreement_score":0.0063195843,"about_ca_system_score_codex":0.00068462826,"about_ca_system_score_gemma":0.0021196404,"threshold_uncertainty_score":0.024833977},"labels":[],"label_agreement":null},{"id":"W3178956269","doi":"10.1007/s10664-021-10004-6","title":"Evaluating the impact of falsely detected performance bug-inducing changes in JIT models","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal; Concordia University","funders":"","keywords":"Leverage (statistics); Software bug; Computer science; Software; Software quality; Empirical research; Software quality assurance; Quality (philosophy); Capability Maturity Model; Source code; Software engineering; Software development; Operating system; Artificial intelligence","score_opus":0.09454756759896381,"score_gpt":0.3660511802990926,"score_spread":0.2715036127001288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3178956269","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99078673,0.00033163384,0.0067620757,0.00029367878,0.000079418714,0.000032979995,0.000305711,0.0004710425,0.00093670975],"genre_scores_gemma":[0.99392796,0.000047624133,0.0052637956,0.000053583874,0.000019087434,0.0000118137295,0.00038800802,0.00006602788,0.0002221672],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98631895,0.0069007315,0.0009060206,0.0022665448,0.002918123,0.0006895311],"domain_scores_gemma":[0.5707826,0.383166,0.016538428,0.018204713,0.009001589,0.0023067065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015667267,0.0010076228,0.0006448692,0.0014643052,0.00063574995,0.0015359956,0.001591823,0.0020867425,0.0014371319],"category_scores_gemma":[0.1943422,0.0005431952,0.0010639754,0.0008711187,0.0013059058,0.0023500277,0.001165968,0.0021726678,0.000258705],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006643844,0.0027005475,0.47332674,0.0007251075,0.001526198,0.0006788447,0.0007320732,0.37685367,0.0149976425,0.0038857672,0.003050113,0.114879444],"study_design_scores_gemma":[0.00026046872,0.0031243344,0.10073688,0.00010146585,0.0008157821,0.00040667277,0.00034821924,0.8775497,0.0120962,0.0036918267,0.00078738574,0.000081086255],"about_ca_topic_score_codex":0.0069970135,"about_ca_topic_score_gemma":0.011818459,"teacher_disagreement_score":0.015667267,"about_ca_system_score_codex":0.0012678988,"about_ca_system_score_gemma":0.0018331595,"threshold_uncertainty_score":0.08285743},"labels":[],"label_agreement":null},{"id":"W3180512116","doi":"10.1007/s10664-021-09969-1","title":"The secret life of test smells - an empirical study on test smell evolution and maintenance","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code smell; Test (biology); Code refactoring; Empirical research; Computer science; Reliability engineering; Engineering; Software; Software quality; Software development; Statistics","score_opus":0.02457635823879043,"score_gpt":0.29435137768838804,"score_spread":0.2697750194495976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3180512116","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.998623,0.00015588295,0.00039297898,0.00007426805,0.0000026004252,0.0000052373966,0.000060296756,0.000011228198,0.00067444495],"genre_scores_gemma":[0.99955696,0.000031143227,0.00012783123,0.000011847089,0.0000037980435,0.0000032200578,0.00008725509,0.000006584931,0.00017136421],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9974942,0.00079902285,0.00025905643,0.00021905715,0.0010425512,0.0001860485],"domain_scores_gemma":[0.8514193,0.08641453,0.04217948,0.0071879686,0.006651992,0.006146668],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0047808522,0.00018416664,0.00024103232,0.0023676832,0.00071941095,0.0015299442,0.0007630025,0.00082968245,0.0016562941],"category_scores_gemma":[0.065046534,0.00021438388,0.0004089064,0.0018878623,0.0014472421,0.004139775,0.001699617,0.0013061739,0.00032737668],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030094106,0.00045261512,0.96782255,0.00006115201,0.00006195097,0.00035286398,0.005216712,0.00055330206,0.001419418,0.0014746253,0.0003333651,0.021950584],"study_design_scores_gemma":[0.000013418273,0.0004279023,0.98616064,0.000053194693,0.00003535335,0.000662176,0.0043838364,0.0040376415,0.0010552513,0.0021536273,0.0009839991,0.000033053253],"about_ca_topic_score_codex":0.0018281317,"about_ca_topic_score_gemma":0.002583406,"teacher_disagreement_score":0.9952192,"about_ca_system_score_codex":0.00082396145,"about_ca_system_score_gemma":0.00056980253,"threshold_uncertainty_score":0.025283873},"labels":[],"label_agreement":null},{"id":"W3183042663","doi":"10.1109/icsme52107.2021.00036","title":"Design Smells in Deep Learning Programs: An Empirical Study","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code smell; Computer science; Context (archaeology); Artificial intelligence; Artificial neural network; Relevance (law); Empirical research; Software engineering; Software; Quality (philosophy); Machine learning; Software quality; Software development","score_opus":0.08701109926562463,"score_gpt":0.3627289947378883,"score_spread":0.27571789547226366,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3183042663","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99796695,0.00027052834,0.00093589595,0.00011764788,0.0000035759256,0.00004457751,0.00007914196,0.000039579943,0.0005420785],"genre_scores_gemma":[0.99711657,0.00027372953,0.001721446,0.00008652974,0.0000058425317,0.00006613909,0.00026421068,0.000031215062,0.0004343499],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98863703,0.0035188962,0.0013637965,0.0009262365,0.004957326,0.00059672433],"domain_scores_gemma":[0.77522945,0.15353285,0.044539016,0.0058760867,0.017424023,0.0033986347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013203866,0.0005322062,0.0004211911,0.0027581025,0.0006771868,0.0016390645,0.0011572273,0.0012880609,0.0016296267],"category_scores_gemma":[0.10259164,0.00063197996,0.00059469807,0.0023521977,0.0015202638,0.0036681707,0.0018949718,0.0022067367,0.00046039632],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005581701,0.002089809,0.8831021,0.000979245,0.000117894735,0.0015265715,0.024399696,0.0022709405,0.0025509247,0.0007006942,0.0022523727,0.079451635],"study_design_scores_gemma":[0.0000894875,0.0025448212,0.9129463,0.0010873266,0.00015361769,0.0027855157,0.033206325,0.02913373,0.0058000647,0.0012022965,0.010895391,0.00015512077],"about_ca_topic_score_codex":0.0019369214,"about_ca_topic_score_gemma":0.0031504764,"teacher_disagreement_score":0.013203866,"about_ca_system_score_codex":0.0013736659,"about_ca_system_score_gemma":0.0010766955,"threshold_uncertainty_score":0.06982952},"labels":[],"label_agreement":null},{"id":"W3183521984","doi":"10.1145/3468744.3468754","title":"Reflections on","year":2021,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Readability; Computer science; Code (set theory); Style (visual arts); Software engineering; Programming language; World Wide Web; Art; Visual arts","score_opus":0.04293661375924791,"score_gpt":0.3177162513290067,"score_spread":0.2747796375697588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3183521984","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019950592,0.0068065096,0.00037224017,0.90198624,0.043354254,0.000032707798,0.00029162513,0.000116972165,0.045044396],"genre_scores_gemma":[0.03216395,0.0064437143,0.0005592123,0.7305245,0.010099635,0.000109161854,0.00027153263,0.00037467512,0.21945368],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9892461,0.0040594777,0.00032238106,0.0014600301,0.0025962463,0.0023156449],"domain_scores_gemma":[0.97461593,0.006259148,0.0007442493,0.0014364633,0.009171398,0.0077728163],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.008193395,0.0010530555,0.00061335845,0.0012698851,0.013686001,0.013633976,0.0032752485,0.014787888,0.06580309],"category_scores_gemma":[0.03775458,0.0005314367,0.00094037247,0.0011512015,0.0061598774,0.008969685,0.008200245,0.03080027,0.021621708],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027627302,0.000019818064,0.00019809308,0.000049848848,0.000002870461,0.00014961866,0.0031595107,0.000012404005,0.000099134944,0.0040081576,0.9850825,0.007190423],"study_design_scores_gemma":[0.000004073838,0.0000107007345,0.0003196476,0.00016583416,0.0000023566774,0.000070117625,0.007214831,0.0000071811123,0.00011583275,0.00060663983,0.99147063,0.000012209042],"about_ca_topic_score_codex":0.050556544,"about_ca_topic_score_gemma":0.08237867,"teacher_disagreement_score":0.9341969,"about_ca_system_score_codex":0.014366952,"about_ca_system_score_gemma":0.009144029,"threshold_uncertainty_score":0.22013324},"labels":[],"label_agreement":null},{"id":"W3184068421","doi":"10.1142/s0218194021500327","title":"An Extensible Compiler for Implementing Software Design Patterns as Concise Language Constructs","year":2021,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software design pattern; Computer science; Structural pattern; Design pattern; Programming language; Compiler; Software design; Software engineering; Software; Software development","score_opus":0.019229357488186768,"score_gpt":0.29895102840481125,"score_spread":0.2797216709166245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3184068421","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004689316,0.00008912447,0.97765064,0.00011825284,0.00008453907,0.00017366318,0.00014598657,0.015700545,0.0013478966],"genre_scores_gemma":[0.020311758,0.00017070677,0.97330767,0.00014844308,0.00003079735,0.0002573567,0.0005395249,0.0030304606,0.0022032366],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983596,0.00030240996,0.00035528466,0.00023784877,0.0006195919,0.00012523272],"domain_scores_gemma":[0.9954782,0.0012465973,0.0006378369,0.0016677409,0.0008162316,0.00015337882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002824952,0.0010032634,0.00051921874,0.0013224869,0.00046015935,0.0018849354,0.0018599498,0.0009866693,0.0018545329],"category_scores_gemma":[0.008325828,0.0013775956,0.0014273898,0.0010743113,0.0007666795,0.003058487,0.0019742015,0.0026272247,0.0013926107],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045752808,0.0006643175,0.008396788,0.001613542,0.00025309328,0.0019132526,0.0030363824,0.02686758,0.11408313,0.1735406,0.03135796,0.6378158],"study_design_scores_gemma":[0.00043132764,0.0006270085,0.0029378834,0.0007835899,0.00045461883,0.004231145,0.00038322338,0.30511132,0.13099812,0.09002971,0.46366835,0.0003436716],"about_ca_topic_score_codex":0.0009191092,"about_ca_topic_score_gemma":0.0015823118,"teacher_disagreement_score":0.002824952,"about_ca_system_score_codex":0.0005026414,"about_ca_system_score_gemma":0.0023711482,"threshold_uncertainty_score":0.014939904},"labels":[],"label_agreement":null},{"id":"W3185065804","doi":"10.1145/3452379","title":"PLIERS","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Computer-Human Interaction","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Carnegie Mellon University; U.S. Department of Defense; National Science Foundation","keywords":"Computer science; Immutability; Programming language; Usability; Java; Process (computing); Software engineering; Human–computer interaction","score_opus":0.029425097163547555,"score_gpt":0.3058972013839144,"score_spread":0.27647210422036683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3185065804","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012847507,0.0005683539,0.78724664,0.0019326605,0.0003931753,0.0013520604,0.0021275845,0.04849024,0.1450418],"genre_scores_gemma":[0.115563616,0.00093587686,0.67507863,0.0017484097,0.00019733475,0.0016052314,0.00589142,0.016775293,0.18220426],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99607724,0.0011822003,0.00036078325,0.00075801904,0.0013995784,0.00022214897],"domain_scores_gemma":[0.9909672,0.0036915457,0.00045721515,0.0022434453,0.002377762,0.00026285736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004963714,0.0011057667,0.00033474038,0.0012762923,0.00096215395,0.0033204153,0.0015853401,0.00081460027,0.071338445],"category_scores_gemma":[0.014561291,0.00092754315,0.0007680133,0.0006000017,0.0012708906,0.005455641,0.0029689316,0.0024210417,0.024901506],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003832829,0.00016742019,0.0024756424,0.00132933,0.00005335284,0.00032078387,0.004754725,0.0022475864,0.020798283,0.18629846,0.10193387,0.67923725],"study_design_scores_gemma":[0.00007625461,0.00019532612,0.00087329716,0.00024844694,0.0000313059,0.00044644251,0.0005595123,0.006486796,0.016421072,0.03250813,0.942103,0.000050400206],"about_ca_topic_score_codex":0.0011015534,"about_ca_topic_score_gemma":0.0018662731,"teacher_disagreement_score":0.071338445,"about_ca_system_score_codex":0.00093346275,"about_ca_system_score_gemma":0.001695926,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3185295980","doi":"10.1007/s10664-021-10099-x","title":"Clones in deep learning code: what, where, and why?","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Deep learning; Computer science; Artificial intelligence; Timeline; Source lines of code; Context (archaeology); Python (programming language); Dependability; Software development; Java; Software evolution; Software system; Software engineering; Machine learning; Software quality; Code (set theory); Software; Programming language; Software construction; Biology","score_opus":0.017447204121601945,"score_gpt":0.26496860410912637,"score_spread":0.24752139998752443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3185295980","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9512247,0.0013385043,0.039841104,0.0036812627,0.000065420296,0.000028881075,0.0001635042,0.00034169524,0.003315028],"genre_scores_gemma":[0.9914062,0.00022959086,0.006951903,0.00022584996,0.000030777475,0.000017398514,0.000107738306,0.00011067392,0.0009199992],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9934296,0.002333936,0.0003049168,0.0011222424,0.0022838723,0.0005254023],"domain_scores_gemma":[0.8946272,0.072831996,0.010017982,0.012021314,0.008804229,0.0016973104],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062619396,0.0003959183,0.0005534513,0.0013968645,0.001075719,0.002350294,0.0011174148,0.0022338745,0.0024517518],"category_scores_gemma":[0.12980126,0.0005521412,0.0004591212,0.001920214,0.0055271336,0.010591605,0.0023267716,0.0031851744,0.00038329104],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042183712,0.000358155,0.5429724,0.00034046732,0.00019705975,0.0005441908,0.0075585553,0.022031022,0.0045865704,0.09009612,0.0066610095,0.3242327],"study_design_scores_gemma":[0.0001375919,0.00055900315,0.19112568,0.00088022626,0.0003808433,0.0020739737,0.0090963645,0.27042732,0.019516613,0.49325302,0.012379593,0.00016973504],"about_ca_topic_score_codex":0.0063897497,"about_ca_topic_score_gemma":0.009067275,"teacher_disagreement_score":0.0063897497,"about_ca_system_score_codex":0.002328058,"about_ca_system_score_gemma":0.0021302428,"threshold_uncertainty_score":0.0331167},"labels":[],"label_agreement":null},{"id":"W3188152090","doi":"10.1109/ms.2021.3101249","title":"From GWT to Angular: An Experiment Report on Migrating a Legacy Web Application","year":2021,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Berger (Canada)","funders":"","keywords":"Computer science; Web application; World Wide Web; Software; Operating system","score_opus":0.020332023076295298,"score_gpt":0.3055672975441192,"score_spread":0.28523527446782393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3188152090","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9770583,0.00021705819,0.008315713,0.0004296539,0.0003046217,0.00068659644,0.0008771403,0.006757742,0.0053532855],"genre_scores_gemma":[0.94039094,0.0003003897,0.045063637,0.0005136007,0.00006781356,0.0007178586,0.0034798593,0.0014263925,0.008039359],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99719185,0.0008988494,0.00028489958,0.00053496676,0.00072679954,0.0003625999],"domain_scores_gemma":[0.99162877,0.0031576827,0.00038039425,0.0022882645,0.0015711974,0.0009737057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030214651,0.00093701814,0.00043697478,0.00067729945,0.0010729183,0.00094228645,0.0014368786,0.0009767495,0.0023544442],"category_scores_gemma":[0.011758434,0.00041436986,0.00045986622,0.0007701526,0.0009041275,0.0018796095,0.0016506212,0.0016252362,0.0014692799],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012629717,0.027333457,0.084450066,0.0042330488,0.00070321717,0.006230399,0.017250953,0.018912164,0.29498085,0.0040467083,0.054942213,0.47428727],"study_design_scores_gemma":[0.0058024116,0.049246296,0.17733303,0.00048398707,0.0014380685,0.0048371577,0.01696872,0.16583446,0.41621143,0.005445585,0.15569445,0.0007043596],"about_ca_topic_score_codex":0.0059813876,"about_ca_topic_score_gemma":0.005549867,"teacher_disagreement_score":0.0059813876,"about_ca_system_score_codex":0.0003977153,"about_ca_system_score_gemma":0.001034691,"threshold_uncertainty_score":0.01597923},"labels":[],"label_agreement":null},{"id":"W3188230742","doi":"10.1109/rew53955.2021.00024","title":"Issue Link Label Recovery and Prediction for Open Source Software","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Task (project management); Software; Open source; Open source software; Focus (optics); Link (geometry); Machine learning; Data science; Selection (genetic algorithm); Feature selection; Artificial intelligence; Software engineering; Data mining; Systems engineering; Engineering","score_opus":0.03474133248225599,"score_gpt":0.2999174149204027,"score_spread":0.2651760824381467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3188230742","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6703168,0.0015862325,0.30826116,0.0012836488,0.00027256392,0.00026642103,0.0030121894,0.012295183,0.0027057326],"genre_scores_gemma":[0.7979588,0.00037633473,0.18662573,0.0001302072,0.00017074168,0.00020734988,0.01145652,0.00056992116,0.0025043837],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99317765,0.0025416154,0.00048559392,0.0014934142,0.0018786829,0.00042305936],"domain_scores_gemma":[0.9299383,0.04462734,0.0093383705,0.007235683,0.0077570863,0.0011031147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009657288,0.0014802896,0.0007796127,0.011858121,0.0016995327,0.0023368238,0.0017977692,0.0024310504,0.0010059179],"category_scores_gemma":[0.04696382,0.00047948572,0.0008890387,0.0052961167,0.00089260284,0.004415826,0.0023791092,0.003975792,0.0013631716],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004431087,0.0017098814,0.3163039,0.00061890547,0.00022544985,0.0006300496,0.0017774842,0.119273804,0.00813385,0.004006255,0.024072444,0.5228048],"study_design_scores_gemma":[0.000021320411,0.000119779594,0.027141407,0.00007567731,0.000055730532,0.0001939927,0.00048012284,0.9471084,0.011079641,0.0081752185,0.0054964554,0.00005223898],"about_ca_topic_score_codex":0.0058558397,"about_ca_topic_score_gemma":0.008002286,"teacher_disagreement_score":0.011858121,"about_ca_system_score_codex":0.0011889102,"about_ca_system_score_gemma":0.0014198052,"threshold_uncertainty_score":0.051073253},"labels":[],"label_agreement":null},{"id":"W3189144731","doi":"10.1109/icces51350.2021.9489247","title":"Retracted: An Adaptable and Extensible Code Smell Detection Approach","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":true,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Code smell; Code refactoring; Mistake; Computer science; Code (set theory); Product (mathematics); Source code; Extensibility; Software engineering; Programming language; Software quality; Software; Software development","score_opus":0.028708291884656494,"score_gpt":0.2535352720499648,"score_spread":0.2248269801653083,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3189144731","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034403738,0.00065453176,0.912446,0.00090395316,0.00021023382,0.0009711371,0.00062664045,0.042559996,0.0072237323],"genre_scores_gemma":[0.27706435,0.00045365802,0.6988999,0.00083567895,0.00013146843,0.0004614314,0.0028242306,0.003125335,0.01620388],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9939366,0.0009668085,0.00047824607,0.0013536104,0.0029025702,0.00036218637],"domain_scores_gemma":[0.9898482,0.0018618251,0.001276813,0.0036201798,0.0028529312,0.00053990626],"candidate_categories":["metaresearch","research_integrity"],"consensus_categories":[],"category_scores_codex":[0.005076257,0.001535907,0.0009994416,0.003895202,0.00092426135,0.002540786,0.005169901,0.0017436077,0.003317937],"category_scores_gemma":[0.019517902,0.00089085253,0.0017344491,0.0015543862,0.001177814,0.0053146826,0.0055775046,0.0025752566,0.0019369943],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042131162,0.0006361822,0.026128089,0.00047497923,0.0001692504,0.0018152359,0.0027258391,0.01604473,0.017850256,0.013250638,0.016432622,0.9040508],"study_design_scores_gemma":[0.00009643269,0.0008152051,0.014931266,0.00044215686,0.00038471937,0.002579943,0.0017955563,0.8194533,0.028223062,0.03606138,0.09492436,0.0002926771],"about_ca_topic_score_codex":0.0058831265,"about_ca_topic_score_gemma":0.008676011,"teacher_disagreement_score":0.9982564,"about_ca_system_score_codex":0.00090293644,"about_ca_system_score_gemma":0.0022566787,"threshold_uncertainty_score":0.02684617},"labels":[],"label_agreement":null},{"id":"W3189287397","doi":"10.1007/978-3-030-75855-4_14","title":"Cost Optimization of Software Quality Assurance","year":2021,"lang":"en","type":"book-chapter","venue":"Studies in big data","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sheridan College","funders":"","keywords":"Quality assurance; Software quality assurance; Computer science; Software quality; Software; Quality costs; Reliability engineering; Business; Engineering; Risk analysis (engineering); Software development; Operations management; Cost control; Operating system","score_opus":0.3398517357818084,"score_gpt":0.4067914611748718,"score_spread":0.06693972539306342,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3189287397","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019875446,0.018227998,0.83986795,0.0025826057,0.0004607864,0.00007690152,0.00016135744,0.00044369383,0.118303195],"genre_scores_gemma":[0.602689,0.01730503,0.28608125,0.0004533498,0.00089556864,0.0002476958,0.00030774897,0.0012222878,0.09079808],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99863523,0.0004511302,0.00003316313,0.000110934205,0.00066727586,0.00010222663],"domain_scores_gemma":[0.9978948,0.0014891939,0.00010994348,0.000185615,0.0002745908,0.000045911707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001092942,0.0010738432,0.0009314775,0.001181273,0.00036511268,0.0015148686,0.0011925539,0.00077376125,0.0076708314],"category_scores_gemma":[0.0059997695,0.00047759252,0.000650602,0.0017951091,0.0009589466,0.002006923,0.0007834642,0.0015509141,0.0008733619],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000071466646,0.00007347769,0.00034750882,0.00028671024,0.000040780356,0.000028774608,0.00006584136,0.26598927,0.0020740512,0.37424237,0.013599184,0.34318066],"study_design_scores_gemma":[0.000018815133,0.00010499225,0.0012515673,0.00015637356,0.00004529626,0.00008025896,0.000063618856,0.6198935,0.0020928758,0.34682012,0.029439613,0.000032931708],"about_ca_topic_score_codex":0.002559323,"about_ca_topic_score_gemma":0.0021964523,"teacher_disagreement_score":0.0076708314,"about_ca_system_score_codex":0.0026170735,"about_ca_system_score_gemma":0.0009278432,"threshold_uncertainty_score":0.025661469},"labels":[],"label_agreement":null},{"id":"W3190371720","doi":"10.1007/978-981-16-1927-4_15","title":"Sometimes, Cloning Is a Sound Design Decision!","year":2021,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code refactoring; Cloning (programming); Code (set theory); Source code; Computer science; Software engineering; clone (Java method); Programming language; Software; Biology; Genetics","score_opus":0.0504611684080111,"score_gpt":0.2800818788749672,"score_spread":0.22962071046695612,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3190371720","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004503632,0.008360905,0.6487384,0.05132886,0.009147499,0.00015827405,0.00032189675,0.004317315,0.2731233],"genre_scores_gemma":[0.091025956,0.008251992,0.43357348,0.026265802,0.0025055045,0.0003365294,0.00039296728,0.0067650946,0.4308827],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949929,0.0016732479,0.00023533414,0.00074435706,0.0020287894,0.00032536994],"domain_scores_gemma":[0.98887014,0.005937992,0.00033469952,0.0031605973,0.0013639159,0.0003327622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006472765,0.0015133173,0.0011334821,0.001064057,0.0028737425,0.008061324,0.0030383507,0.0038040802,0.033169553],"category_scores_gemma":[0.021884859,0.0011389946,0.0010312466,0.0013205808,0.009184703,0.01881449,0.0035818238,0.008836974,0.020971837],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000042107356,0.000034271845,0.00020640614,0.00023772795,0.000020686106,0.00011705024,0.000766708,0.00044841584,0.0015982958,0.8368062,0.07868731,0.08103483],"study_design_scores_gemma":[0.00001643922,0.00004992678,0.00009374431,0.0002774489,0.000035233767,0.0005686531,0.00039136535,0.0010791834,0.0031570035,0.5099285,0.4843625,0.000039919603],"about_ca_topic_score_codex":0.0013952026,"about_ca_topic_score_gemma":0.0025408368,"teacher_disagreement_score":0.033169553,"about_ca_system_score_codex":0.0017412754,"about_ca_system_score_gemma":0.0019910973,"threshold_uncertainty_score":0.110963166},"labels":[],"label_agreement":null},{"id":"W3191590060","doi":"10.1109/icjece.2021.3084850","title":"Prediction of Software Effort in the Early Stage of Software Development: A Hybrid Model","year":2021,"lang":"en","type":"article","venue":"Canadian Journal of Electrical and Computer Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Software; Perceptron; Software development; Predictive modelling; Baseline (sea); Multilayer perceptron; Machine learning; Key (lock); Data mining; Artificial intelligence; Artificial neural network","score_opus":0.014507290066262874,"score_gpt":0.19440652307570072,"score_spread":0.17989923300943786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3191590060","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6292393,0.0020334448,0.35805938,0.0017196833,0.000118972814,0.000109021596,0.0017409271,0.0018650579,0.0051141637],"genre_scores_gemma":[0.9716573,0.00038349876,0.024358483,0.00008763853,0.00003906101,0.0000733441,0.00067269267,0.000030641804,0.0026973982],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994129,0.00016190077,0.000038677685,0.00019805251,0.00010973805,0.00007868051],"domain_scores_gemma":[0.99867743,0.0007717261,0.00015830487,0.00007415596,0.00026196017,0.000056374865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001462624,0.00072263385,0.0005877348,0.0011033041,0.0002112036,0.001072532,0.0014060335,0.0010659355,0.0011451418],"category_scores_gemma":[0.003026931,0.00031556416,0.0006238602,0.0011115146,0.00026262237,0.0013662096,0.0006893024,0.00085912086,0.0004614275],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035296014,0.00037378471,0.028351435,0.00015033716,0.00016292416,0.0002058494,0.00015586875,0.8591284,0.0019285005,0.0015424637,0.002780555,0.104866914],"study_design_scores_gemma":[0.000004580651,0.000045085835,0.0028556227,0.000009906897,0.000016371281,0.000016775934,0.000012089015,0.99601245,0.00019436731,0.0006221356,0.00020488266,0.0000058146184],"about_ca_topic_score_codex":0.009246527,"about_ca_topic_score_gemma":0.009192256,"teacher_disagreement_score":0.009246527,"about_ca_system_score_codex":0.00079287373,"about_ca_system_score_gemma":0.0006744045,"threshold_uncertainty_score":0.01838541},"labels":[],"label_agreement":null},{"id":"W3192457683","doi":"10.1007/s10664-021-09980-6","title":"An empirical study of same-day releases of popular packages in the npm ecosystem","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Alberta; Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Reuse; Schedule; Exploratory research; Software engineering; Operating system; Engineering","score_opus":0.030237698177021738,"score_gpt":0.3166212142887117,"score_spread":0.28638351611169,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3192457683","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99908435,0.000021968011,0.00006909941,0.00005060438,0.0000016691937,0.0000062125036,0.00007301123,0.000005373885,0.0006876795],"genre_scores_gemma":[0.99861956,0.000033167416,0.00024915143,0.000035938432,0.000007264933,0.000012953367,0.00030244948,0.000009188328,0.00073039863],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9979109,0.0007876703,0.00011874274,0.000320137,0.00057508826,0.00028750638],"domain_scores_gemma":[0.94000894,0.032847553,0.013727369,0.0033966608,0.0056047128,0.0044147316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031788987,0.00023750436,0.00030485707,0.0020154219,0.0013684161,0.0019552566,0.0009649773,0.0011830695,0.0026129917],"category_scores_gemma":[0.033366907,0.0003274307,0.00030046506,0.002458965,0.0014762221,0.0039806003,0.0016173882,0.0019518937,0.000719413],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043151423,0.0020423615,0.97087437,0.00006253525,0.000059663646,0.000529864,0.008188143,0.00034215307,0.0014927596,0.0011495486,0.001125786,0.013701227],"study_design_scores_gemma":[0.000016431726,0.00026226754,0.98731893,0.000017701448,0.000018040906,0.000251652,0.008226448,0.0022071237,0.00027955,0.00024108765,0.0011405336,0.000020327157],"about_ca_topic_score_codex":0.011871191,"about_ca_topic_score_gemma":0.024066865,"teacher_disagreement_score":0.011871191,"about_ca_system_score_codex":0.0013823682,"about_ca_system_score_gemma":0.00081731926,"threshold_uncertainty_score":0.023604155},"labels":[],"label_agreement":null},{"id":"W3194111178","doi":"10.1007/s00766-021-00360-6","title":"A validation of QDAcity-RE for domain modeling using qualitative data analysis","year":2021,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Deutsche Forschungsgemeinschaft; Friedrich-Alexander-Universität Erlangen-Nürnberg","keywords":"Traceability; Domain (mathematical analysis); Computer science; Set (abstract data type); Data mining; Artificial intelligence; Machine learning; Mathematics; Software engineering; Programming language","score_opus":0.22658466325772372,"score_gpt":0.42070231335018665,"score_spread":0.19411765009246293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3194111178","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12879649,0.00011960834,0.8509747,0.0007912534,0.00014687922,0.006021189,0.0011949495,0.0025144136,0.009440564],"genre_scores_gemma":[0.304248,0.0000481163,0.6883318,0.00018730843,0.000010287443,0.0052536284,0.0008025351,0.00026905283,0.0008493368],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.8712289,0.093355104,0.007075565,0.00786217,0.019413175,0.0010651551],"domain_scores_gemma":[0.48082668,0.34962112,0.011112144,0.091941975,0.064763136,0.0017349671],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.14415862,0.0013123545,0.00094499346,0.004229848,0.0018127504,0.0043551144,0.0032889596,0.0016755412,0.0041559003],"category_scores_gemma":[0.2854567,0.0008385969,0.0012064627,0.0024258776,0.0036631182,0.0042150556,0.0054915347,0.0022242873,0.0009850388],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002499286,0.005192743,0.049723372,0.0057789907,0.00069792283,0.0008527463,0.04729341,0.06372481,0.07742477,0.09738431,0.0069964984,0.64243114],"study_design_scores_gemma":[0.0014845356,0.004003456,0.04270984,0.0034974318,0.00030882526,0.0011566042,0.024914816,0.6619432,0.140765,0.05634819,0.06214371,0.0007244384],"about_ca_topic_score_codex":0.003673242,"about_ca_topic_score_gemma":0.0028755357,"teacher_disagreement_score":0.14415862,"about_ca_system_score_codex":0.0037081565,"about_ca_system_score_gemma":0.0053882003,"threshold_uncertainty_score":0.76239276},"labels":[],"label_agreement":null},{"id":"W3194114348","doi":"10.1007/s10664-021-09986-0","title":"Empirical evaluation of tools for hairy requirements engineering tasks","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Empirical research; Systems engineering; Software engineering; Engineering; Mathematics","score_opus":0.13428343853465402,"score_gpt":0.37559869310442306,"score_spread":0.24131525456976904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3194114348","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9883222,0.00016914672,0.00805306,0.000110920344,0.000022855633,0.00037983354,0.00015714044,0.00047951762,0.002305364],"genre_scores_gemma":[0.97550327,0.00015146339,0.021923466,0.00007815606,0.000014789796,0.0004066977,0.00057671434,0.00014684997,0.0011985728],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9840821,0.009330714,0.0013851408,0.0009633723,0.0036604754,0.00057821436],"domain_scores_gemma":[0.6508252,0.30127162,0.011887866,0.016908186,0.0154944975,0.003612624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016138071,0.0008280085,0.00048132226,0.002902193,0.0007698184,0.001945292,0.0019958273,0.0013307058,0.0030329246],"category_scores_gemma":[0.17741063,0.0004990667,0.00052271114,0.0015848774,0.0010305081,0.0027361743,0.0029392568,0.001194169,0.0009403788],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013083866,0.022270802,0.09134758,0.0043714345,0.000404185,0.0010657054,0.025707634,0.020671995,0.035250485,0.0047894316,0.0065940944,0.7744428],"study_design_scores_gemma":[0.0084390445,0.06642298,0.5210679,0.002815084,0.0011117549,0.0019191714,0.028499957,0.25890565,0.06279161,0.00939551,0.038017448,0.0006139723],"about_ca_topic_score_codex":0.0019227756,"about_ca_topic_score_gemma":0.0031549537,"teacher_disagreement_score":0.016138071,"about_ca_system_score_codex":0.0012926715,"about_ca_system_score_gemma":0.0015704982,"threshold_uncertainty_score":0.085347354},"labels":[],"label_agreement":null},{"id":"W3194283088","doi":"10.1145/3479497","title":"The \"Shut the f**k up\" Phenomenon: Characterizing Incivility in Open Source Code Review Discussions","year":2021,"lang":"en","type":"article","venue":"Proceedings of the ACM on Human-Computer Interaction","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Incivility; Civility; Context (archaeology); Phenomenon; Internet privacy; Computer science; Computer security; World Wide Web; Public relations; Social psychology; Psychology; Political science; Law; Epistemology","score_opus":0.07765418840429965,"score_gpt":0.3631479555744238,"score_spread":0.28549376717012415,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3194283088","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9772141,0.00069466676,0.00880613,0.0036601794,0.000082712315,0.00017609655,0.00008015438,0.00012269989,0.009163167],"genre_scores_gemma":[0.99629104,0.00023563825,0.0017489335,0.00075268844,0.000049484872,0.00019744552,0.000059638594,0.000057540805,0.0006075924],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9437389,0.04084769,0.0028512427,0.0027260003,0.0072679375,0.0025683157],"domain_scores_gemma":[0.7082104,0.22088422,0.04893406,0.0073175337,0.0089665875,0.005687233],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.024629055,0.0005827453,0.0006648565,0.006262058,0.008354625,0.007484009,0.002219041,0.0035469886,0.0022389786],"category_scores_gemma":[0.14682332,0.00075233413,0.00064792845,0.0032368754,0.013552961,0.0096386215,0.013548449,0.003888013,0.00041406025],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009203492,0.00006454361,0.03659263,0.00028989292,0.000029004046,0.00077826285,0.93917406,0.00010028756,0.0014938073,0.0056574945,0.0011257716,0.014602126],"study_design_scores_gemma":[0.000023166025,0.00012019665,0.046176527,0.0007055226,0.000035029563,0.0016088476,0.9167534,0.0012352968,0.0010991794,0.008906314,0.023209346,0.0001272429],"about_ca_topic_score_codex":0.002925303,"about_ca_topic_score_gemma":0.0030168549,"teacher_disagreement_score":0.992516,"about_ca_system_score_codex":0.00526225,"about_ca_system_score_gemma":0.0037134096,"threshold_uncertainty_score":0.13025248},"labels":[],"label_agreement":null},{"id":"W3194585484","doi":"10.1145/3468264.3468609","title":"How disabled tests manifest in test maintainability challenges?","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Technical debt; Software maintenance; Computer science; Test (biology); Maintainability; Software engineering; Code smell; Quality (philosophy); Empirical research; Reliability engineering; Risk analysis (engineering); Java; Codebase; Software quality; Software; Software development; Engineering","score_opus":0.02653790395737383,"score_gpt":0.26622074050630595,"score_spread":0.23968283654893213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3194585484","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9004146,0.008794112,0.048408195,0.021436425,0.0003443067,0.00014819777,0.0023097445,0.0016518938,0.016492406],"genre_scores_gemma":[0.9877072,0.00081657525,0.007660063,0.0013419939,0.0001478581,0.000064280386,0.00086709374,0.00033167552,0.0010631888],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9744187,0.0073305005,0.0021148294,0.00475291,0.009688552,0.0016944776],"domain_scores_gemma":[0.6520679,0.2350342,0.058295846,0.029011404,0.020710668,0.004880084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017505735,0.000652355,0.00060577336,0.0071033463,0.0012214901,0.004394209,0.0024058118,0.0022192302,0.003358806],"category_scores_gemma":[0.2390302,0.0007408134,0.0006449238,0.0051255133,0.003971436,0.010999033,0.0028026123,0.002585562,0.00090873253],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026566617,0.0001982597,0.6644576,0.0005522746,0.00014387254,0.0017437566,0.01176557,0.0019072553,0.0026577592,0.012991079,0.0076890914,0.2956278],"study_design_scores_gemma":[0.00007318792,0.0005896784,0.77287495,0.0020245605,0.00035825957,0.016563985,0.020589897,0.019619353,0.00799169,0.08867096,0.07039421,0.00024920644],"about_ca_topic_score_codex":0.005499741,"about_ca_topic_score_gemma":0.005261938,"teacher_disagreement_score":0.017505735,"about_ca_system_score_codex":0.0019243831,"about_ca_system_score_gemma":0.0021685853,"threshold_uncertainty_score":0.09258026},"labels":[],"label_agreement":null},{"id":"W3195323992","doi":"10.1145/3468264.3473132","title":"Term interrelations and trends in software engineering","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Term (time); Software; Embedding; Word (group theory); Software development; Rapid prototyping; Test (biology)","score_opus":0.015068468563573653,"score_gpt":0.25972209075109454,"score_spread":0.24465362218752087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3195323992","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84495115,0.018364757,0.113217086,0.00235492,0.00045555204,0.00012993797,0.005957706,0.0007411862,0.013827691],"genre_scores_gemma":[0.9231531,0.00537545,0.062208503,0.00014114012,0.00016924638,0.00017035744,0.0058455113,0.00023229896,0.0027043526],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9973471,0.000825149,0.0004977417,0.00057059596,0.00064147986,0.00011797298],"domain_scores_gemma":[0.98070645,0.01325302,0.0026109666,0.0010817372,0.0020871942,0.00026072102],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0020727275,0.0004409754,0.00033500887,0.01087986,0.00073533383,0.0024485956,0.00048011597,0.00068506127,0.0020574615],"category_scores_gemma":[0.022012444,0.0003058442,0.0005383116,0.016793653,0.0010302186,0.0068495297,0.0017574187,0.0010611588,0.00069268857],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060240575,0.00018643694,0.21226317,0.0027600762,0.0003532591,0.0010297602,0.01016257,0.014609289,0.01606136,0.047492016,0.009943455,0.68453616],"study_design_scores_gemma":[0.00011008224,0.0008176364,0.4312849,0.0023058455,0.00077320467,0.0056103887,0.0153810615,0.13422185,0.018598596,0.20020677,0.19037122,0.00031842056],"about_ca_topic_score_codex":0.0022321332,"about_ca_topic_score_gemma":0.005078716,"teacher_disagreement_score":0.9891201,"about_ca_system_score_codex":0.0007386066,"about_ca_system_score_gemma":0.00078795233,"threshold_uncertainty_score":0.010961711},"labels":[],"label_agreement":null},{"id":"W3195370809","doi":"","title":"A Systematic Literature Review of Automated Query Reformulations in Source Code Search.","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Information retrieval; Table (database); Software; Source code; Code (set theory); World Wide Web; Data mining; Programming language; Set (abstract data type)","score_opus":0.046844861937170766,"score_gpt":0.22376366864668265,"score_spread":0.17691880670951188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3195370809","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029072443,0.98692316,0.0020185146,0.001201265,0.00021479651,0.0023874165,0.0024732726,0.000078742756,0.0017956507],"genre_scores_gemma":[0.02850088,0.9445784,0.012899956,0.0023006,0.00012205649,0.007614967,0.0034024296,0.000069709706,0.00051106606],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.96280867,0.015187156,0.012405259,0.0020791031,0.006917567,0.00060226914],"domain_scores_gemma":[0.83871126,0.12561513,0.01341236,0.0031292231,0.0182376,0.0008943513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.041375037,0.002015628,0.0054160953,0.042858407,0.0015298852,0.0040389886,0.0044564707,0.0024858413,0.0067897937],"category_scores_gemma":[0.1616191,0.0014537935,0.0051696384,0.03957432,0.0021004113,0.007833536,0.0045874943,0.0017908324,0.0012659095],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012237352,0.000030313442,0.0007589524,0.89803445,0.0013300127,0.0001543217,0.0015070343,0.00013160564,0.00033846192,0.00069490937,0.004877529,0.092020005],"study_design_scores_gemma":[0.00011850343,0.00019065395,0.0027675934,0.93792856,0.009359385,0.00036940168,0.001793492,0.00016356733,0.00043552963,0.0007903582,0.046026584,0.000056392204],"about_ca_topic_score_codex":0.01303672,"about_ca_topic_score_gemma":0.034755927,"teacher_disagreement_score":0.042858407,"about_ca_system_score_codex":0.008683827,"about_ca_system_score_gemma":0.039919917,"threshold_uncertainty_score":0.21881473},"labels":[],"label_agreement":null},{"id":"W3195404945","doi":"","title":"Software Batch Testing to Reduce Build Test Executions","year":2020,"lang":"en","type":"dissertation","venue":"Spectrum Research Repository (Concordia University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Concordia University","keywords":"Computer science; Commit; Test (biology); Test case; Bisection method; Code coverage; Reliability engineering; Software; Programming language; Algorithm; Database; Machine learning; Engineering","score_opus":0.03419390070201119,"score_gpt":0.28700472231156815,"score_spread":0.252810821609557,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3195404945","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.159417,0.0014147186,0.7481036,0.001030334,0.00043955067,0.0015382936,0.0009834456,0.06539305,0.021679953],"genre_scores_gemma":[0.51893324,0.00029701626,0.46308926,0.0005805162,0.00009677913,0.00097478763,0.0020478421,0.0038765867,0.0101039205],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9902952,0.00246426,0.00071510026,0.001470969,0.004211556,0.00084300473],"domain_scores_gemma":[0.94885224,0.022562245,0.0029762997,0.0159815,0.008189303,0.0014384501],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054455292,0.0022222372,0.0011675433,0.0020044204,0.0006953428,0.0018808239,0.005642893,0.00094755046,0.012680778],"category_scores_gemma":[0.031069575,0.0011565611,0.0014026586,0.0014341753,0.0011786311,0.0038519406,0.002427668,0.0030056087,0.0036419304],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023950601,0.0015247533,0.01578219,0.0010216773,0.00026263917,0.0005310942,0.00059220364,0.18095353,0.08661143,0.019434994,0.03304819,0.6578422],"study_design_scores_gemma":[0.000486955,0.0018158951,0.008203559,0.0001562173,0.00020165282,0.00050817325,0.0002148543,0.84881175,0.096128695,0.0148812765,0.028439779,0.00015114165],"about_ca_topic_score_codex":0.007134651,"about_ca_topic_score_gemma":0.009167036,"teacher_disagreement_score":0.012680778,"about_ca_system_score_codex":0.0018724129,"about_ca_system_score_gemma":0.00402841,"threshold_uncertainty_score":0.0424214},"labels":[],"label_agreement":null},{"id":"W3195642814","doi":"10.1145/3475960.3475984","title":"Heterogeneous ensemble imputation for software development effort estimation","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Imputation (statistics); Computer science; Decision tree; Missing data; Data mining; Regression; Support vector machine; Artificial intelligence; Statistics; Machine learning; Mathematics","score_opus":0.019758071731206885,"score_gpt":0.2785844139466412,"score_spread":0.2588263422154343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3195642814","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036094,0.0007617013,0.9606873,0.00017425118,0.00005207771,0.000058106914,0.00040454383,0.0008082474,0.00095983635],"genre_scores_gemma":[0.6749853,0.00060867576,0.32091522,0.00009966637,0.00009061686,0.000186218,0.0020299247,0.00010959873,0.0009747583],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951199,0.0024092086,0.000289817,0.0009211309,0.0010300962,0.00022993895],"domain_scores_gemma":[0.9868896,0.0076198746,0.0012147807,0.0018101741,0.0022328782,0.00023261644],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007886636,0.0009057347,0.0017043777,0.0031596061,0.00049866183,0.0010725944,0.0013526633,0.0007564844,0.0009951702],"category_scores_gemma":[0.02264645,0.0004502085,0.0015473629,0.0034863125,0.0002455106,0.0016520654,0.0010322994,0.0015468274,0.0004600153],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001797095,0.00016960251,0.022572853,0.00018408071,0.00070358545,0.00012966711,0.00018811441,0.5796646,0.0013164309,0.0042827507,0.0028378475,0.38777074],"study_design_scores_gemma":[0.000009042534,0.0000681827,0.003691892,0.00004286958,0.00006669363,0.00004001543,0.000045168676,0.98608,0.0013651971,0.0073020975,0.0012662658,0.000022562346],"about_ca_topic_score_codex":0.0048846477,"about_ca_topic_score_gemma":0.0055115265,"teacher_disagreement_score":0.007886636,"about_ca_system_score_codex":0.00065845984,"about_ca_system_score_gemma":0.0010519646,"threshold_uncertainty_score":0.041709006},"labels":[],"label_agreement":null},{"id":"W3195653486","doi":"10.48550/arxiv.2108.07474","title":"A grounded theory of Community Package Maintenance\\n Organizations-Registered Report","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure","funders":"","keywords":"Grounded theory; Plan (archaeology); Extant taxon; Context (archaeology); Function (biology); Computer science; Knowledge management; Process management; Qualitative research; Management science; Business; Engineering; Sociology; Geography","score_opus":0.08843218740740404,"score_gpt":0.21526413801410474,"score_spread":0.12683195060670072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3195653486","genre_codex":"methods","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13377635,0.0013755184,0.50737005,0.036220655,0.00048448797,0.007840042,0.0008959832,0.00026049544,0.31177646],"genre_scores_gemma":[0.78286767,0.00061790546,0.20401289,0.0016959076,0.000026611451,0.0038984395,0.000385772,0.00006864974,0.0064261374],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9546939,0.03898902,0.00087613094,0.0018658494,0.0026706713,0.00090445235],"domain_scores_gemma":[0.9626464,0.02976439,0.0012840193,0.0025952996,0.002395293,0.0013145693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.033873625,0.0006604884,0.00062168785,0.0034747056,0.0070039514,0.009250411,0.00360646,0.0024512527,0.005158319],"category_scores_gemma":[0.022718001,0.0006666239,0.0007240494,0.0040063392,0.035996,0.012605735,0.006885838,0.004199992,0.0006224583],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000023904164,0.00013671345,0.0012034852,0.00037482483,0.000009561688,0.0001708796,0.15628721,0.00085262797,0.00022236923,0.8206678,0.0026515678,0.017399011],"study_design_scores_gemma":[0.00011808221,0.00012280692,0.0014572552,0.0025672603,0.000025266983,0.00024067894,0.2800318,0.0067852377,0.00070733944,0.5731306,0.13476217,0.00005148451],"about_ca_topic_score_codex":0.009693364,"about_ca_topic_score_gemma":0.011562951,"teacher_disagreement_score":0.033873625,"about_ca_system_score_codex":0.018758792,"about_ca_system_score_gemma":0.018149134,"threshold_uncertainty_score":0.17914301},"labels":[],"label_agreement":null},{"id":"W3196038764","doi":"","title":"Creating Better Action Plans for Writing Tasks via Vocabulary-Based Planning","year":2018,"lang":"en","type":"article","venue":"Conference on Computer Supported Cooperative Work","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Vocabulary; Task (project management); Action (physics); Context (archaeology); Plan (archaeology); Set (abstract data type); Action plan; Domain (mathematical analysis); Human–computer interaction; Task analysis; Automation; Knowledge management; Artificial intelligence; Process management; Engineering","score_opus":0.0633750200978596,"score_gpt":0.3246366454533156,"score_spread":0.261261625355456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196038764","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22835244,0.00031879038,0.7408358,0.00070624554,0.00009765834,0.0031038376,0.00087636977,0.005215379,0.020493438],"genre_scores_gemma":[0.38899758,0.00020144013,0.60354584,0.00009541183,0.000019318573,0.001765822,0.0012157098,0.00037785046,0.0037810013],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9940942,0.0036730778,0.0004747725,0.0007939396,0.00075841736,0.0002055915],"domain_scores_gemma":[0.9557351,0.032557063,0.0032892106,0.004892303,0.0027160575,0.00081035506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074385754,0.0013921049,0.00048171193,0.001977521,0.0011488972,0.0029822334,0.0016559149,0.0010020045,0.006517295],"category_scores_gemma":[0.051143378,0.00054070767,0.0008700544,0.0011677492,0.0016374398,0.004508455,0.0025804588,0.0014810086,0.002301276],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007533178,0.0011937902,0.029285066,0.0031054746,0.00010949273,0.0010852799,0.12717408,0.026386645,0.04624658,0.035891403,0.011687369,0.7170815],"study_design_scores_gemma":[0.0011470955,0.0032157223,0.045334265,0.0028900723,0.000568675,0.00285871,0.13123296,0.3422352,0.05912189,0.16170767,0.24883586,0.00085188745],"about_ca_topic_score_codex":0.0024509856,"about_ca_topic_score_gemma":0.0042203246,"teacher_disagreement_score":0.0074385754,"about_ca_system_score_codex":0.0010356081,"about_ca_system_score_gemma":0.0029779903,"threshold_uncertainty_score":0.039339364},"labels":[],"label_agreement":null},{"id":"W3196126762","doi":"10.1109/tse.2021.3106247","title":"Dependency Smells in JavaScript Projects","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code smell; Computer science; Dependency (UML); Commit; JavaScript; Code refactoring; Dependency graph; Software engineering; Security bug; World Wide Web; Popularity; Computer security; Data science; Software; Software development; Software quality; Software security assurance; Database; Programming language; Information security","score_opus":0.018662997445303465,"score_gpt":0.2381131029342694,"score_spread":0.21945010548896593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196126762","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9960983,0.00051651296,0.0014868863,0.00025794472,0.000008509169,0.00002781513,0.00078092044,0.00010687999,0.00071623264],"genre_scores_gemma":[0.99444276,0.00039085638,0.0023251106,0.00008296178,0.000020141419,0.00007472847,0.0021043916,0.00006604085,0.00049306307],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9888214,0.0030488337,0.0020746617,0.0016773775,0.0037677935,0.0006099677],"domain_scores_gemma":[0.7712713,0.109123446,0.09120293,0.0092389835,0.015140022,0.0040232986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074820286,0.00029626087,0.00034841752,0.007023984,0.00097427913,0.001586357,0.0006663695,0.00085343193,0.0007717127],"category_scores_gemma":[0.081374176,0.00044928587,0.0004488257,0.007600724,0.0009128446,0.0034782272,0.003124486,0.0012932043,0.00031354977],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009694454,0.000063779575,0.95849115,0.0002665282,0.00004612022,0.00048340842,0.009732634,0.00037511947,0.0012595925,0.00022711867,0.0016944885,0.02726311],"study_design_scores_gemma":[0.000006017791,0.00010358118,0.984209,0.0001699406,0.000023158926,0.0010725698,0.0074362475,0.0017140328,0.00074659433,0.00042817896,0.004045311,0.000045308738],"about_ca_topic_score_codex":0.0046733376,"about_ca_topic_score_gemma":0.009186126,"teacher_disagreement_score":0.0074820286,"about_ca_system_score_codex":0.00094832486,"about_ca_system_score_gemma":0.00085656927,"threshold_uncertainty_score":0.0395692},"labels":[],"label_agreement":null},{"id":"W3196571700","doi":"10.18293/seke2021-163","title":"Evaluating a Tool for Creating Bug Report Assignment Recommenders (S)","year":2021,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Workload; Field (mathematics); Software bug; Triage; Software; Software engineering; World Wide Web; Operating system","score_opus":0.05676680508889584,"score_gpt":0.3214043356662734,"score_spread":0.2646375305773776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196571700","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89103925,0.00057397294,0.065208316,0.0004619594,0.00022053765,0.0028047352,0.0019558435,0.034382444,0.0033529028],"genre_scores_gemma":[0.67211163,0.000325052,0.31354567,0.00036658774,0.00007564576,0.0015135105,0.0052989577,0.0008502539,0.0059127146],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9914282,0.003950012,0.0010742852,0.0013467133,0.0018004334,0.00040030808],"domain_scores_gemma":[0.9152226,0.06664792,0.0026451123,0.005442719,0.0071259583,0.0029157652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013550494,0.0014368917,0.0010462945,0.0027688912,0.00083660474,0.0019151461,0.0023446598,0.0025116487,0.003269094],"category_scores_gemma":[0.059548352,0.00090653426,0.0008336643,0.0013574777,0.0004228139,0.0025775724,0.0017132947,0.0014350496,0.002233068],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007500227,0.014024368,0.06572462,0.0034047456,0.0007050173,0.0008442077,0.0063814213,0.011744251,0.03172453,0.000967135,0.019538106,0.8374413],"study_design_scores_gemma":[0.007576285,0.049990393,0.16879992,0.0008757257,0.0018426287,0.0024989464,0.006113703,0.62553173,0.06469209,0.0018679966,0.06908965,0.0011209337],"about_ca_topic_score_codex":0.004461557,"about_ca_topic_score_gemma":0.0071788426,"teacher_disagreement_score":0.013550494,"about_ca_system_score_codex":0.00067520555,"about_ca_system_score_gemma":0.0013895187,"threshold_uncertainty_score":0.071662724},"labels":[],"label_agreement":null},{"id":"W3196823214","doi":"10.48550/arxiv.2108.13587","title":"T3-Vis: a visual analytic framework for Training and fine-Tuning Transformers in NLP","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Suite; Transformer; Computer science; Visualization; Artificial intelligence; Architecture; Focus (optics); Process (computing); Machine learning; Natural language processing; Programming language; Engineering","score_opus":0.09381849016606994,"score_gpt":0.24494936643288998,"score_spread":0.15113087626682004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196823214","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00096635584,0.00008226375,0.989599,0.0002629143,0.000025145077,0.00006297079,0.0003199446,0.0077134334,0.0009679898],"genre_scores_gemma":[0.056045655,0.0002924556,0.9392195,0.00022991406,0.000045082026,0.00035164072,0.00078778807,0.0019372758,0.001090682],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981425,0.00092479127,0.00012662786,0.00034845067,0.00036998477,0.000087719076],"domain_scores_gemma":[0.9923381,0.0058307196,0.0003144202,0.0007258712,0.0005249157,0.00026603742],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047652074,0.002358303,0.0009205925,0.0030826938,0.0010183569,0.0046854126,0.0031785017,0.0020752016,0.017981494],"category_scores_gemma":[0.024646822,0.0009778421,0.002297289,0.0013634722,0.0023876512,0.006133902,0.0044915536,0.0036468427,0.003728343],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008228464,0.00020534974,0.0022005942,0.0016575588,0.00020091122,0.00057302054,0.0030211085,0.2221968,0.016276224,0.24901085,0.044910587,0.45892417],"study_design_scores_gemma":[0.00007798691,0.000072241615,0.00019146131,0.00020867791,0.00003368935,0.00014474851,0.0002545563,0.72282016,0.006018367,0.23924856,0.030878698,0.000050795374],"about_ca_topic_score_codex":0.0063617313,"about_ca_topic_score_gemma":0.009461291,"teacher_disagreement_score":0.017981494,"about_ca_system_score_codex":0.0016626765,"about_ca_system_score_gemma":0.0022661919,"threshold_uncertainty_score":0.06015408},"labels":[],"label_agreement":null},{"id":"W3196940930","doi":"10.48550/arxiv.2109.00659","title":"Semantic Slicing of Architectural Change Commits: Towards Semantic Design Review","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Commit; Software engineering; Architectural pattern; Slicing; Process (computing); Java; Programming language; Software; Database; Software design; Software development; World Wide Web","score_opus":0.1612324899678375,"score_gpt":0.23301951233541876,"score_spread":0.07178702236758125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196940930","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020691171,0.00068348757,0.9632973,0.0006756645,0.00011420687,0.0005234023,0.0013692102,0.010397727,0.002247889],"genre_scores_gemma":[0.09225461,0.0005803989,0.89783394,0.00021209498,0.00009526888,0.0004203574,0.004875475,0.0020642763,0.0016635925],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9753686,0.008205354,0.0029620037,0.003084488,0.00972863,0.0006507708],"domain_scores_gemma":[0.92251986,0.027616547,0.01108877,0.020033235,0.017546004,0.0011955969],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016263105,0.0016266435,0.0013108368,0.0122777205,0.0013060072,0.005292611,0.002298334,0.001847659,0.0025790061],"category_scores_gemma":[0.0769996,0.0011202615,0.0019650944,0.005030726,0.0016445443,0.008130098,0.005167724,0.002243376,0.001834341],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027801018,0.00029953287,0.01750273,0.0020870569,0.00018580291,0.0007999962,0.007874173,0.013288301,0.027884841,0.045456782,0.021286042,0.8630568],"study_design_scores_gemma":[0.00014687161,0.000553537,0.015021308,0.0021561084,0.00041793007,0.002004133,0.0062930253,0.3804548,0.122057766,0.16279916,0.3077043,0.00039097504],"about_ca_topic_score_codex":0.002549126,"about_ca_topic_score_gemma":0.0035900034,"teacher_disagreement_score":0.016263105,"about_ca_system_score_codex":0.0019334513,"about_ca_system_score_gemma":0.0062648836,"threshold_uncertainty_score":0.08600861},"labels":[],"label_agreement":null},{"id":"W3196992266","doi":"10.1002/smr.2378","title":"A study of refactorings during software change tasks","year":2021,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Norges Forskningsråd","keywords":"Code refactoring; Workflow; Computer science; Software engineering; Scope (computer science); Software; Task (project management); Adaptation (eye); Software development; Programming language; Systems engineering; Engineering; Database","score_opus":0.03438675384689678,"score_gpt":0.297320505641714,"score_spread":0.2629337517948172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196992266","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9972249,0.00007229857,0.0019648527,0.000069983624,0.00000393945,0.00009221985,0.000022758328,0.000021545406,0.00052746595],"genre_scores_gemma":[0.9960006,0.0000700709,0.0033462541,0.00004723763,0.0000049843206,0.00009408249,0.00005751394,0.000019094885,0.0003601005],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9780523,0.013361461,0.0016906972,0.0016863863,0.0040317574,0.0011772942],"domain_scores_gemma":[0.75583625,0.17193696,0.03533212,0.010780242,0.021421969,0.004692386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018653901,0.0006690341,0.000606734,0.0026546852,0.002473623,0.0024565894,0.0018886314,0.0017771858,0.0009133366],"category_scores_gemma":[0.16437452,0.001073428,0.00037285165,0.0013125014,0.0023103477,0.0024149593,0.0025024144,0.0020503805,0.00030788686],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003695206,0.0022262249,0.3473232,0.0004381329,0.00007162002,0.0016492363,0.5751572,0.0006260314,0.013220596,0.0007696647,0.0005642486,0.05758428],"study_design_scores_gemma":[0.00018272262,0.0033834714,0.6332438,0.00041955884,0.00008274064,0.002029687,0.32273486,0.008492486,0.012468808,0.0023289581,0.014266312,0.00036666426],"about_ca_topic_score_codex":0.006520371,"about_ca_topic_score_gemma":0.009674368,"teacher_disagreement_score":0.018653901,"about_ca_system_score_codex":0.0023031784,"about_ca_system_score_gemma":0.0024643547,"threshold_uncertainty_score":0.09865242},"labels":[],"label_agreement":null},{"id":"W3197075544","doi":"10.17771/pucrio.acad.48681","title":"CATALOGING DEPENDENCY INJECTION ANTI-PATTERNS IN SOFTWARE SYSTEMS","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Inversa Systems (Canada)","funders":"Pontifícia Universidade Católica do Rio de Janeiro; Coordenação de Aperfeiçoamento de Pessoal de Nível Superior","keywords":"Cataloging; Dependency (UML); Computer science; Software engineering; Information retrieval; World Wide Web","score_opus":0.0346676210012855,"score_gpt":0.2775621449887255,"score_spread":0.24289452398744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197075544","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7871066,0.011081207,0.18519764,0.0013931498,0.0001303176,0.0009712888,0.003309217,0.0019563688,0.00885431],"genre_scores_gemma":[0.81379706,0.004985255,0.17219368,0.0003444261,0.00011593022,0.00069072534,0.005623071,0.0002582985,0.0019915192],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9866166,0.0025686768,0.0034747883,0.0020694474,0.004689314,0.0005812137],"domain_scores_gemma":[0.90026075,0.04801051,0.024334976,0.0106499,0.015329444,0.0014144548],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062900954,0.00066937663,0.0008201161,0.016039664,0.0012459031,0.0027015142,0.001256806,0.0013129959,0.0011416794],"category_scores_gemma":[0.035988357,0.0006972575,0.0007182909,0.015305098,0.0013499594,0.00411211,0.001735516,0.00082029,0.0003903584],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034116255,0.0004565695,0.38195032,0.00491564,0.00020045933,0.002037407,0.008336814,0.0048032734,0.012850554,0.015551057,0.005095031,0.56346166],"study_design_scores_gemma":[0.00014196568,0.0020463665,0.57951784,0.006186609,0.00077654794,0.024737377,0.015595457,0.05688514,0.03893185,0.070094176,0.20462725,0.00045936648],"about_ca_topic_score_codex":0.0029781146,"about_ca_topic_score_gemma":0.004185693,"teacher_disagreement_score":0.016039664,"about_ca_system_score_codex":0.0013255512,"about_ca_system_score_gemma":0.0022246893,"threshold_uncertainty_score":0.03326559},"labels":[],"label_agreement":null},{"id":"W3197188256","doi":"10.1016/j.infsof.2021.106756","title":"Early prediction for merged vs abandoned code changes in modern code reviews","year":2021,"lang":"en","type":"preprint","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Code review; Code (set theory); Machine learning; Classifier (UML); Artificial intelligence; Source code; Suite; Software; Source lines of code; Process (computing); Empirical research; Software inspection; Software engineering; Software development; Software quality; Programming language; Statistics","score_opus":0.02435593368266787,"score_gpt":0.27361690913263004,"score_spread":0.24926097544996217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197188256","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9498756,0.004013439,0.027969427,0.0016793845,0.0007028214,0.00016516665,0.0040821037,0.0018783828,0.009633609],"genre_scores_gemma":[0.98827606,0.00030276418,0.00683868,0.00014925266,0.00018775476,0.000037054833,0.0020606623,0.00020062165,0.0019470694],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9926662,0.0011675996,0.00062062044,0.0015283723,0.0032414244,0.0007756897],"domain_scores_gemma":[0.80723184,0.13428606,0.020033425,0.00892669,0.024887973,0.004634019],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005784745,0.00052540173,0.0006068309,0.005787952,0.0011038681,0.0026764404,0.00094124733,0.0018084919,0.0035515942],"category_scores_gemma":[0.100118876,0.00050645054,0.0008190485,0.0024683634,0.00076079095,0.003658534,0.0012544659,0.0023954657,0.0015696208],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026621192,0.00023830074,0.7893838,0.0008001067,0.00030647774,0.0012679475,0.0009358509,0.0056069666,0.009717418,0.007186622,0.018339535,0.16355498],"study_design_scores_gemma":[0.00017888786,0.0008709149,0.7199842,0.0006468838,0.0006076207,0.00379912,0.0013416258,0.19418325,0.027232755,0.022351854,0.028589984,0.00021297784],"about_ca_topic_score_codex":0.002570335,"about_ca_topic_score_gemma":0.005038762,"teacher_disagreement_score":0.005787952,"about_ca_system_score_codex":0.00096876756,"about_ca_system_score_gemma":0.0015650108,"threshold_uncertainty_score":0.030593038},"labels":[],"label_agreement":null},{"id":"W3197531090","doi":"","title":"A Data Extraction Algorithm From Open Source Software Project Repositories for Building Duration Estimation Models: Case Study of GitHub","year":2020,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Duration (music); Computer science; Estimation; Schedule; Software; Data extraction; Variable (mathematics); Work (physics); Data mining; Data science; Engineering; Systems engineering","score_opus":0.07751349942747296,"score_gpt":0.3566497924798842,"score_spread":0.2791362930524113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197531090","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15993707,0.0004600357,0.8291007,0.00060169806,0.000052751726,0.0006034415,0.0042334422,0.0030892922,0.0019215585],"genre_scores_gemma":[0.2362129,0.00035879598,0.75364685,0.00004858239,0.000027350776,0.00069397583,0.0077493903,0.0001495627,0.0011126132],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99806553,0.0006106543,0.00032792782,0.00044929542,0.0004491549,0.00009741839],"domain_scores_gemma":[0.9879067,0.0072814557,0.0011833681,0.0012957073,0.0021030316,0.00022987703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036903273,0.0008180329,0.0006149978,0.003911467,0.0006973413,0.0014044256,0.0010359224,0.000884788,0.0008360274],"category_scores_gemma":[0.018091338,0.00041803127,0.0010041193,0.004903295,0.00033591807,0.0022068447,0.0013131609,0.0010343203,0.00082514307],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025447874,0.000649824,0.08738345,0.0007033557,0.00016037379,0.00065634877,0.0016324708,0.07316828,0.007975934,0.0068824063,0.008192487,0.8123406],"study_design_scores_gemma":[0.00006129564,0.00030663874,0.041137993,0.00023751495,0.00013597049,0.00072683103,0.0014827409,0.9071767,0.015027291,0.009392053,0.024199855,0.00011521061],"about_ca_topic_score_codex":0.008509718,"about_ca_topic_score_gemma":0.01223664,"teacher_disagreement_score":0.008509718,"about_ca_system_score_codex":0.0007860197,"about_ca_system_score_gemma":0.0022246726,"threshold_uncertainty_score":0.019516587},"labels":[],"label_agreement":null},{"id":"W3199062392","doi":"","title":"A comparative exploration of FreeBSD bug lifetimes","year":2010,"lang":"en","type":"article","venue":"Singapore Management University Institutional Knowledge (InK) (Singapore Management University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Eclipse; Computer science; Software bug; Data mining; Programming language; Software","score_opus":0.026512186530792653,"score_gpt":0.23982817801459438,"score_spread":0.21331599148380173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3199062392","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9743974,0.0033099726,0.010945538,0.0002994227,0.000030338897,0.000035010587,0.0066792793,0.0012536083,0.0030494363],"genre_scores_gemma":[0.9795517,0.0005828208,0.007825124,0.00003182508,0.00001801718,0.000024158722,0.011050422,0.00017494564,0.00074107095],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9976157,0.0006026297,0.00022859177,0.000411132,0.0009943788,0.00014768033],"domain_scores_gemma":[0.9513391,0.03627954,0.003975351,0.0026603632,0.004967208,0.0007784803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004453288,0.0005411658,0.00063999,0.014995557,0.00039953098,0.00126338,0.0007341775,0.00046421768,0.0013782785],"category_scores_gemma":[0.029568363,0.00021209962,0.0008039262,0.006409658,0.00036339863,0.0028247125,0.0010035062,0.00051949127,0.00045516467],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012933859,0.00028655707,0.5607716,0.001091385,0.0006451049,0.00035562244,0.002052717,0.037199423,0.0038589311,0.002935044,0.006624131,0.38288608],"study_design_scores_gemma":[0.00008761812,0.0016455549,0.6911416,0.00032515344,0.00033262337,0.0021610123,0.002389265,0.26428235,0.008969071,0.008185278,0.020305896,0.00017469135],"about_ca_topic_score_codex":0.0035815274,"about_ca_topic_score_gemma":0.0062234118,"teacher_disagreement_score":0.014995557,"about_ca_system_score_codex":0.0006110208,"about_ca_system_score_gemma":0.00045315825,"threshold_uncertainty_score":0.023551524},"labels":[],"label_agreement":null},{"id":"W3199935553","doi":"","title":"Measuring API documentation on the Web","year":2011,"lang":"en","type":"article","venue":"Singapore Management University Institutional Knowledge (InK) (Singapore Management University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Documentation; World Wide Web; Software documentation; Computer science; Social media; Social software; Software; Software development; Social web; User analysis; Software development process; Human–computer interaction","score_opus":0.050546420966246126,"score_gpt":0.20796008772622185,"score_spread":0.15741366675997573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3199935553","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99078697,0.0003633727,0.000689241,0.00009093739,0.000010292314,0.00003844972,0.00092024915,0.00010479215,0.0069956775],"genre_scores_gemma":[0.9925585,0.00039067367,0.0024315652,0.00003438692,0.000034886827,0.00007164267,0.0025221505,0.000049707276,0.0019065354],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9940797,0.0015354551,0.0010391664,0.0004679706,0.0026011846,0.00027653377],"domain_scores_gemma":[0.92123264,0.03973717,0.019437514,0.003228955,0.014533666,0.0018300944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038016832,0.00027878585,0.00040672292,0.012784909,0.00054831145,0.0026818821,0.0004197307,0.0007437135,0.0020095215],"category_scores_gemma":[0.043953784,0.00025528081,0.00029640866,0.008479505,0.00038611918,0.0037119298,0.0014741094,0.00058761914,0.0011738156],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014285753,0.00021691379,0.9048335,0.0003891252,0.0001173658,0.0001936004,0.0042644623,0.00058153336,0.0020589512,0.00047011752,0.0024123471,0.08431926],"study_design_scores_gemma":[0.000007982324,0.00016569512,0.9814052,0.00013896,0.000080572725,0.0004940496,0.0045956,0.0038526002,0.0019878305,0.00030452406,0.006928589,0.00003845217],"about_ca_topic_score_codex":0.0027957615,"about_ca_topic_score_gemma":0.0032195952,"teacher_disagreement_score":0.012784909,"about_ca_system_score_codex":0.000504603,"about_ca_system_score_gemma":0.00050670066,"threshold_uncertainty_score":0.020105481},"labels":[],"label_agreement":null},{"id":"W3200568639","doi":"10.1007/s10664-021-10025-1","title":"A study of how Docker Compose is used to compose multi-component systems","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Component (thermodynamics); Computer science; Container (type theory); Software deployment; Operating system; Component-based software engineering; Software; Leverage (statistics); Web application; Software engineering; AKA; Database; World Wide Web; Software system; Engineering","score_opus":0.07160242864086612,"score_gpt":0.3200402561237651,"score_spread":0.24843782748289894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3200568639","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9844694,0.000048532034,0.008762855,0.00010753305,0.000009030151,0.000039150444,0.000020827889,0.00015556399,0.0063871257],"genre_scores_gemma":[0.9843068,0.000056217435,0.011941672,0.000036836595,0.0000049151536,0.000016487294,0.000072413255,0.00010822346,0.0034564196],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9980824,0.0009085393,0.0000970707,0.00021114896,0.0005715563,0.00012924412],"domain_scores_gemma":[0.9688757,0.024654223,0.001448322,0.002756404,0.0017696572,0.0004956226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032428196,0.00038636624,0.00033318633,0.0012402341,0.0015881654,0.002044312,0.00083639077,0.0009310951,0.0035446247],"category_scores_gemma":[0.03268315,0.00039132318,0.0003573215,0.0014980414,0.0013103655,0.0034623654,0.0010345707,0.0011280978,0.00047845364],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017378045,0.0065450165,0.31472546,0.00075664197,0.00028603434,0.0031858492,0.110269345,0.03303064,0.0545902,0.053149592,0.0060680853,0.4156554],"study_design_scores_gemma":[0.00031662363,0.004083628,0.38680056,0.00028710792,0.00043855116,0.0041059162,0.10189723,0.3354302,0.080495186,0.025371552,0.060454525,0.00031892088],"about_ca_topic_score_codex":0.009062455,"about_ca_topic_score_gemma":0.015047413,"teacher_disagreement_score":0.009062455,"about_ca_system_score_codex":0.0014104847,"about_ca_system_score_gemma":0.0010326564,"threshold_uncertainty_score":0.018019438},"labels":[],"label_agreement":null},{"id":"W3200749310","doi":"10.1002/smr.2395","title":"Behind the scenes: On the relationship between developer experience and refactoring","year":2021,"lang":"en","type":"preprint","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Code refactoring; Technical debt; Computer science; Generalizability theory; Software engineering; Variety (cybernetics); Software development; Software; Programming language; Artificial intelligence; Psychology","score_opus":0.08183480943976264,"score_gpt":0.3301927228445304,"score_spread":0.24835791340476776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3200749310","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9985586,0.00008969981,0.0002760223,0.00011836662,0.000004866573,0.0000058732653,0.000028945364,0.000005545465,0.0009121015],"genre_scores_gemma":[0.9996019,0.000038253635,0.00016310558,0.00001960885,0.00000579663,0.00000560199,0.0000325697,0.0000036544282,0.00012955931],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99339986,0.0035355724,0.00042086485,0.0005266927,0.0015058416,0.00061120413],"domain_scores_gemma":[0.8463386,0.106093705,0.028047564,0.0022144432,0.0097349305,0.007570784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050167493,0.0002541278,0.00025209994,0.0023340352,0.00056236755,0.0017653905,0.0005175459,0.00054730295,0.0022697786],"category_scores_gemma":[0.05256996,0.0002661333,0.00031376508,0.0015156672,0.0012012086,0.0016945573,0.0018564962,0.000941325,0.00024321963],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008658172,0.00010769916,0.973473,0.00004428015,0.000041901938,0.00020254387,0.01582974,0.000079566955,0.00033718717,0.000116532,0.00021695042,0.009463916],"study_design_scores_gemma":[0.000003660498,0.00014246893,0.97639227,0.000042870124,0.000018149498,0.00017511981,0.021845456,0.00052273425,0.00014244842,0.0001049854,0.0005919845,0.000017842503],"about_ca_topic_score_codex":0.0034236442,"about_ca_topic_score_gemma":0.0044130078,"teacher_disagreement_score":0.0050167493,"about_ca_system_score_codex":0.00062773016,"about_ca_system_score_gemma":0.0005878816,"threshold_uncertainty_score":0.026531398},"labels":[],"label_agreement":null},{"id":"W3201809880","doi":"10.3390/e23101274","title":"An Adaptive Rank Aggregation-Based Ensemble Multi-Filter Feature Selection Method in Software Defect Prediction","year":2021,"lang":"en","type":"article","venue":"Entropy","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Filter (signal processing); Rank (graph theory); Feature selection; Computer science; Selection (genetic algorithm); Curse of dimensionality; Data mining; Artificial intelligence; Feature (linguistics); Pattern recognition (psychology); Machine learning; Mathematics","score_opus":0.018183229619683463,"score_gpt":0.28654739203006013,"score_spread":0.2683641624103767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3201809880","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04578069,0.0012506305,0.9504308,0.00023251565,0.0001060905,0.000088824374,0.0003273235,0.0010405114,0.0007425516],"genre_scores_gemma":[0.73218244,0.0008747703,0.2614733,0.00024064661,0.00024821467,0.00032378605,0.0016717939,0.000115839495,0.002869168],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984694,0.00033711258,0.00011852932,0.00038120977,0.0005328235,0.00016094993],"domain_scores_gemma":[0.9980604,0.0007178352,0.00018163162,0.00016963317,0.0008013993,0.00006894596],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022359055,0.0013280881,0.0017174843,0.002519832,0.0005854342,0.00086496456,0.0012259696,0.000905268,0.0010872808],"category_scores_gemma":[0.0047652884,0.00035325045,0.0019148592,0.002087856,0.00032903408,0.0014333772,0.00064746704,0.0010265093,0.00043731442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032376734,0.0002989,0.01458933,0.00019237229,0.00045550565,0.0002341828,0.00015705684,0.20119692,0.012612997,0.0020427266,0.0077064475,0.7601898],"study_design_scores_gemma":[0.000019584206,0.00013865149,0.0036676205,0.000018740224,0.00009329195,0.000112321526,0.000032866727,0.99017227,0.002927856,0.0014496986,0.0013394492,0.000027658882],"about_ca_topic_score_codex":0.0071524736,"about_ca_topic_score_gemma":0.00649509,"teacher_disagreement_score":0.0071524736,"about_ca_system_score_codex":0.00046803415,"about_ca_system_score_gemma":0.0012199159,"threshold_uncertainty_score":0.014221668},"labels":[],"label_agreement":null},{"id":"W3202264790","doi":"10.1007/s10664-021-10016-2","title":"Rotten green tests in Java, Pharo and Python","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Python (programming language); Java; Computer science; Programming language; Unit testing; Test (biology); Software; Biology; Ecology","score_opus":0.02234213422314712,"score_gpt":0.284636052570599,"score_spread":0.26229391834745186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3202264790","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90499026,0.00042547032,0.05048542,0.0007602078,0.00023979954,0.00013034964,0.0027823427,0.001349985,0.038836107],"genre_scores_gemma":[0.9804277,0.000069485795,0.011090555,0.00017001558,0.000056616467,0.00020652187,0.0018337141,0.00079454057,0.0053508803],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98426497,0.009902264,0.0010352904,0.0017412667,0.0020362644,0.0010200508],"domain_scores_gemma":[0.6003603,0.3712649,0.0082448665,0.011724281,0.005660813,0.0027449636],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01407781,0.0007869022,0.0009235585,0.0021881901,0.0011575933,0.0021441765,0.0021470136,0.001102272,0.01667463],"category_scores_gemma":[0.16344479,0.00037010547,0.0017472012,0.0034462004,0.0022645323,0.006842692,0.0030529378,0.002629067,0.0022032284],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012389815,0.0042234124,0.3444031,0.0014348704,0.0021696433,0.00078247744,0.008722796,0.06147906,0.0037093365,0.19967547,0.05490549,0.3061045],"study_design_scores_gemma":[0.0011823855,0.005191617,0.5143338,0.0004995439,0.0010991392,0.00046054213,0.008830825,0.26784265,0.0098702945,0.16028425,0.030072125,0.00033283385],"about_ca_topic_score_codex":0.014387677,"about_ca_topic_score_gemma":0.014646295,"teacher_disagreement_score":0.9859222,"about_ca_system_score_codex":0.0013408873,"about_ca_system_score_gemma":0.0020776817,"threshold_uncertainty_score":0.074451506},"labels":[],"label_agreement":null},{"id":"W3202359992","doi":"10.1145/3475716.3475790","title":"Continuous Software Bug Prediction","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Benchmark (surveying); Computer science; Software bug; Software; Software regression; Set (abstract data type); Data mining; Verification and validation; Software metric; Software development; Software evolution; Software engineering; Machine learning; Software quality; Software construction; Programming language; Engineering","score_opus":0.011991483860232601,"score_gpt":0.23715835986668288,"score_spread":0.22516687600645027,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3202359992","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7013421,0.018700857,0.21097505,0.0024679694,0.00044317552,0.0006540743,0.043688353,0.015429738,0.0062985993],"genre_scores_gemma":[0.8474561,0.002420824,0.09784088,0.00031361688,0.0003147723,0.0004807868,0.04962234,0.00021999977,0.0013306282],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99488586,0.001083972,0.00046606702,0.001848539,0.0014328328,0.0002827713],"domain_scores_gemma":[0.969494,0.015438831,0.0057289232,0.0038574394,0.004507452,0.00097329856],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004028816,0.002005635,0.0011101136,0.005243503,0.0005482111,0.0016555776,0.002681511,0.0014305348,0.0017765083],"category_scores_gemma":[0.02706779,0.00045731352,0.0008870862,0.0052473224,0.0007144998,0.0023924296,0.0015339924,0.0017449893,0.0010609385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079761905,0.0011018765,0.2473201,0.0030067987,0.0005459695,0.00042171677,0.0003360251,0.22227292,0.0057035508,0.003167452,0.046383705,0.4689422],"study_design_scores_gemma":[0.00016302743,0.00096997595,0.08623,0.00042893022,0.00017095135,0.0006556461,0.00017434136,0.8837501,0.0047209086,0.010514511,0.012126366,0.0000952429],"about_ca_topic_score_codex":0.0056895367,"about_ca_topic_score_gemma":0.0045217043,"teacher_disagreement_score":0.0056895367,"about_ca_system_score_codex":0.00080915116,"about_ca_system_score_gemma":0.0014909684,"threshold_uncertainty_score":0.021306634},"labels":[],"label_agreement":null},{"id":"W3203214494","doi":"","title":"Software Functional Sizing Automation from Requirements Written as Triplets","year":2021,"lang":"en","type":"article","venue":"International Conference on Software Engineering Advances","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Automation; Software requirements specification; Software; Functional requirement; Sizing; Software engineering; Software construction; Programming language; Software development; Engineering","score_opus":0.04094549634603421,"score_gpt":0.30153369169442273,"score_spread":0.2605881953483885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3203214494","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03435934,0.000038538088,0.95216477,0.00010087455,0.00004504386,0.00020742379,0.00029060364,0.007296226,0.0054971767],"genre_scores_gemma":[0.34098673,0.00007951033,0.6522516,0.00009470701,0.000027963291,0.00023534252,0.0013809542,0.0018824295,0.0030607576],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99833906,0.00042935667,0.0001719207,0.00023580535,0.0006908472,0.000132948],"domain_scores_gemma":[0.99609,0.002015013,0.00040672818,0.0007918417,0.0006263123,0.00007006148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014690441,0.000983753,0.00065225497,0.0013042259,0.00039090804,0.0011114903,0.0006668138,0.00054371293,0.0061970884],"category_scores_gemma":[0.0059371954,0.0007097986,0.0016092117,0.0008738259,0.00050743046,0.0009995659,0.0011392767,0.0012945716,0.0015530164],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004890519,0.00025599854,0.0025354784,0.0006484579,0.00012277991,0.001295921,0.0005946216,0.23235501,0.12567195,0.061183043,0.009508465,0.5653392],"study_design_scores_gemma":[0.000060310784,0.00024528065,0.00079254434,0.00008979873,0.00005426307,0.0003502356,0.0001280031,0.9058301,0.04750158,0.03542882,0.009487239,0.000031769858],"about_ca_topic_score_codex":0.001001179,"about_ca_topic_score_gemma":0.001853279,"teacher_disagreement_score":0.0061970884,"about_ca_system_score_codex":0.00040457808,"about_ca_system_score_gemma":0.00085512875,"threshold_uncertainty_score":0.02073133},"labels":[],"label_agreement":null},{"id":"W3203221575","doi":"10.3390/a14100289","title":"Comparing Commit Messages and Source Code Metrics for the Prediction Refactoring Activities","year":2021,"lang":"en","type":"article","venue":"Algorithms","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Code refactoring; Computer science; Commit; Java; Code smell; Metric (unit); Source code; Code (set theory); Class (philosophy); Software; Software metric; Data mining; Programming language; Machine learning; Artificial intelligence; Software quality; Software development; Set (abstract data type); Database","score_opus":0.05110975156882941,"score_gpt":0.2871362250593349,"score_spread":0.23602647349050548,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3203221575","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8263209,0.0030825199,0.15234792,0.0005066537,0.0002091005,0.00024066404,0.008193599,0.0068558184,0.0022428718],"genre_scores_gemma":[0.89160246,0.0005732886,0.09043286,0.00005555272,0.0000969315,0.00016076982,0.01551393,0.00028119463,0.0012830045],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99697196,0.0005871729,0.00028334907,0.0008796433,0.0010530903,0.0002247512],"domain_scores_gemma":[0.978779,0.011286748,0.0031759755,0.0019533886,0.0040489133,0.00075595116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031473432,0.0015800993,0.001032269,0.007957292,0.000409162,0.0010641265,0.00086545147,0.0012198293,0.00062596146],"category_scores_gemma":[0.02044155,0.00032052494,0.000671418,0.0045502177,0.00028096375,0.0018587841,0.00066663785,0.0012976874,0.0010793647],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047130312,0.00067691796,0.45374456,0.00049047207,0.0002991487,0.00017003919,0.00035556,0.035281938,0.008451591,0.0006946002,0.009362284,0.49000162],"study_design_scores_gemma":[0.000043160835,0.00057917996,0.21914206,0.00013062786,0.00012379291,0.00035321192,0.0003201749,0.758829,0.01250793,0.0023692492,0.00552767,0.000073882286],"about_ca_topic_score_codex":0.0069764494,"about_ca_topic_score_gemma":0.010480779,"teacher_disagreement_score":0.007957292,"about_ca_system_score_codex":0.0005361067,"about_ca_system_score_gemma":0.0009537348,"threshold_uncertainty_score":0.016644955},"labels":[],"label_agreement":null},{"id":"W3203969037","doi":"10.1145/3470006","title":"Automatic Fault Detection for Deep Learning Programs Using Graph Transformations","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds de recherche du Québec; Institut de Valorisation des Données","keywords":"Computer science; Metamodeling; Graph; Fault detection and isolation; Precision and recall; Construct (python library); Artificial intelligence; Deep learning; Artificial neural network; Machine learning; Software; Process (computing); Data mining; Software engineering; Theoretical computer science; Programming language","score_opus":0.08467335479524384,"score_gpt":0.3304131342028107,"score_spread":0.24573977940756686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3203969037","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09719545,0.00032596182,0.8797896,0.00029794758,0.000028860048,0.00012389805,0.00066146854,0.020601787,0.00097506616],"genre_scores_gemma":[0.66511524,0.00022344186,0.33071285,0.00016353696,0.000013885203,0.00017931582,0.0019639637,0.0006599888,0.0009677524],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9991436,0.00020019844,0.000066122906,0.00026301388,0.0002543865,0.00007269775],"domain_scores_gemma":[0.9977944,0.0012144358,0.00036380647,0.00034766574,0.0002463458,0.00003332853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006114245,0.0013155343,0.00044679042,0.0018897186,0.00027079292,0.0006801109,0.0012528081,0.0007514284,0.0011743071],"category_scores_gemma":[0.004135099,0.0004635462,0.0014656605,0.00063870853,0.0007707631,0.00152025,0.0008397371,0.0011220135,0.00027502596],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031667747,0.00025141297,0.015196503,0.0005430505,0.00014471496,0.00064741273,0.00024580717,0.5558929,0.038669586,0.012777075,0.003667376,0.37164748],"study_design_scores_gemma":[0.000010156953,0.000034169887,0.0004710676,0.00001705417,0.000017265978,0.000053332336,0.000022666523,0.9765786,0.012749745,0.009288052,0.00075221265,0.0000056319723],"about_ca_topic_score_codex":0.006261319,"about_ca_topic_score_gemma":0.009909898,"teacher_disagreement_score":0.006261319,"about_ca_system_score_codex":0.0014061407,"about_ca_system_score_gemma":0.0010896132,"threshold_uncertainty_score":0.012449741},"labels":[],"label_agreement":null},{"id":"W3204270390","doi":"10.1145/3475716.3484487","title":"Semantic Slicing of Architectural Change Commits","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Software engineering; Commit; Slicing; Architectural pattern; Programming language; Process (computing); Java; Software; Database; Software development; Software design; World Wide Web","score_opus":0.0407431961813553,"score_gpt":0.277116203003387,"score_spread":0.23637300682203172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3204270390","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13549297,0.0006780297,0.832366,0.00034111406,0.00012168146,0.00065770716,0.0042252024,0.021496283,0.0046210145],"genre_scores_gemma":[0.36233422,0.0003762426,0.62361515,0.00009620247,0.00005260021,0.00035360872,0.009424963,0.0016461428,0.0021009417],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99569124,0.0007805805,0.0005841001,0.0008164444,0.0018601221,0.00026752736],"domain_scores_gemma":[0.9814464,0.006844138,0.0031247952,0.0040730867,0.004095598,0.0004159378],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003205328,0.0008001949,0.00067538855,0.006495386,0.00079980877,0.0019311366,0.0009937936,0.000755271,0.0020513018],"category_scores_gemma":[0.019273661,0.0005215963,0.0010108935,0.0027052334,0.00081718224,0.0028093276,0.0021770212,0.0008761107,0.0006287181],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054991775,0.00019378558,0.058268506,0.001613076,0.00014362192,0.0015069536,0.0072923866,0.029210662,0.053168274,0.043568864,0.013472845,0.7910111],"study_design_scores_gemma":[0.00008528874,0.0005255221,0.052678213,0.00085582753,0.00026261687,0.0021927515,0.0040986664,0.56631124,0.18134528,0.07338592,0.11801239,0.0002463183],"about_ca_topic_score_codex":0.003784907,"about_ca_topic_score_gemma":0.0057981955,"teacher_disagreement_score":0.006495386,"about_ca_system_score_codex":0.0009738749,"about_ca_system_score_gemma":0.0021284665,"threshold_uncertainty_score":0.01695162},"labels":[],"label_agreement":null},{"id":"W3204491293","doi":"10.1109/tse.2021.3115772","title":"Characterizing and Mitigating Self-Admitted Technical Debt in Build Systems","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Technical debt; Software engineering; Operating system; Software; Software development","score_opus":0.009912218925696303,"score_gpt":0.22976416783745449,"score_spread":0.2198519489117582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3204491293","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97298616,0.0004298349,0.02413975,0.0003070307,0.000017975479,0.00011710937,0.00028781014,0.00045680697,0.0012575954],"genre_scores_gemma":[0.9773691,0.00019494034,0.020488145,0.00009818897,0.000023231112,0.00008537016,0.0008665655,0.00013167498,0.00074275606],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.988746,0.0035318118,0.0013833623,0.0013539225,0.0043203407,0.0006645067],"domain_scores_gemma":[0.89648736,0.053107955,0.029121123,0.005825181,0.0139285335,0.0015298559],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007643808,0.0006584677,0.00048022694,0.006545451,0.001349102,0.0028056577,0.00081694825,0.0011136476,0.0004554142],"category_scores_gemma":[0.07248676,0.0005456097,0.0004571172,0.0029863508,0.0011942988,0.004127734,0.0027876757,0.0011208897,0.00028934507],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022871475,0.00019420395,0.7880252,0.0009881782,0.00009782507,0.0012724922,0.035808932,0.0041975486,0.023274442,0.0015321998,0.0019160134,0.1424642],"study_design_scores_gemma":[0.000022409766,0.00043131434,0.8448782,0.00065268343,0.00014737084,0.002309386,0.024977144,0.0873808,0.017601002,0.0048952335,0.0165105,0.0001939843],"about_ca_topic_score_codex":0.0052863266,"about_ca_topic_score_gemma":0.0102347005,"teacher_disagreement_score":0.007643808,"about_ca_system_score_codex":0.0013353741,"about_ca_system_score_gemma":0.001535728,"threshold_uncertainty_score":0.040424764},"labels":[],"label_agreement":null},{"id":"W3204578495","doi":"10.1145/3470133","title":"A Systematic Review of API Evolution Literature","year":2021,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Software Engineering Research","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Polytechnique Montréal","funders":"","keywords":"Computer science; Artifact (error); Application programming interface; Software evolution; Software engineering; Software; Software development; Compiler; Android (operating system); Data science; World Wide Web; Programming language; Software construction; Artificial intelligence; Operating system","score_opus":0.03773563095177377,"score_gpt":0.3398366848737849,"score_spread":0.3021010539220111,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3204578495","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006575057,0.99582434,0.00046547982,0.00074425974,0.00020121611,0.00016662847,0.00077670475,0.000025885454,0.0011379477],"genre_scores_gemma":[0.0029796266,0.9943321,0.0008642148,0.0005742644,0.00009048117,0.00027668496,0.0006075815,0.000012672992,0.00026234757],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9921943,0.002295111,0.0025487535,0.00072661537,0.001995705,0.00023940616],"domain_scores_gemma":[0.95592946,0.029897837,0.005780379,0.0009281744,0.00692023,0.00054397987],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0087753795,0.0013830207,0.003082105,0.02090438,0.0008565186,0.0029711216,0.0021091397,0.0013763903,0.007248908],"category_scores_gemma":[0.053979263,0.0010175263,0.0040795947,0.02449372,0.0008143069,0.0043132664,0.0018232344,0.0015251173,0.001473343],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011865779,0.000039961793,0.0014266194,0.63723904,0.0014196867,0.00023010871,0.00055373413,0.00021733691,0.00037005614,0.0017137066,0.021084325,0.33558688],"study_design_scores_gemma":[0.00007010857,0.00017166685,0.0065017533,0.7186321,0.0086054765,0.0007990452,0.0005721492,0.000112723246,0.00034720462,0.0015553465,0.26257765,0.000054714867],"about_ca_topic_score_codex":0.008399766,"about_ca_topic_score_gemma":0.026275856,"teacher_disagreement_score":0.97909564,"about_ca_system_score_codex":0.0032574155,"about_ca_system_score_gemma":0.02050913,"threshold_uncertainty_score":0.04640919},"labels":[],"label_agreement":null},{"id":"W3204686189","doi":"10.1007/s10664-021-10032-2","title":"How are project-specific forums utilized? A study of participation, content, and sentiment in the Eclipse ecosystem","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Japan Society for the Promotion of Science","keywords":"Eclipse; Popularity; World Wide Web; Computer science; Bridge (graph theory); Sentiment analysis; Knowledge management; Data science; Public relations; Political science; Artificial intelligence","score_opus":0.09961188966302743,"score_gpt":0.31488797726991474,"score_spread":0.21527608760688732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3204686189","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.998145,0.00002943616,0.0004636102,0.00007167508,0.000004523235,0.000010875023,0.000025024732,0.000007161635,0.0012427428],"genre_scores_gemma":[0.9992894,0.000023101325,0.00029837477,0.00001780747,0.000006234905,0.000013888417,0.000046770216,0.000005973174,0.00029853728],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9960687,0.0024388663,0.00019818523,0.00037231387,0.00052233326,0.00039956666],"domain_scores_gemma":[0.96954465,0.016048491,0.0077972203,0.0011179665,0.0031130386,0.002378726],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006973767,0.00020108696,0.0002571812,0.0021035331,0.0015154429,0.0023807054,0.00031111878,0.00055035006,0.0012981656],"category_scores_gemma":[0.022971151,0.00019822957,0.0002218682,0.0013163672,0.0008317422,0.0027798698,0.0019463735,0.00052935537,0.00020839354],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023500566,0.0001862105,0.85319203,0.00010635799,0.000042628173,0.00040202242,0.10220126,0.0001268336,0.0049561444,0.00089200865,0.0007209975,0.036938358],"study_design_scores_gemma":[0.00001316409,0.0002093494,0.88727355,0.0000853582,0.000033958455,0.00027600225,0.10184477,0.002074606,0.0009705388,0.000684242,0.006500919,0.00003354081],"about_ca_topic_score_codex":0.0016624271,"about_ca_topic_score_gemma":0.0026116674,"teacher_disagreement_score":0.006973767,"about_ca_system_score_codex":0.00076564756,"about_ca_system_score_gemma":0.0005088819,"threshold_uncertainty_score":0.036881268},"labels":[],"label_agreement":null},{"id":"W3204758977","doi":"10.48550/arxiv.2109.14097","title":"How Much Data Analytics is Enough? The ROI of Machine Learning Classification and its Application to Requirements Dependency Classification","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Dependency (UML); Return on investment; Analytics; Artificial intelligence; Selection (genetic algorithm); Region of interest; Software; Machine learning; Data mining","score_opus":0.21359547421863792,"score_gpt":0.2580230380897199,"score_spread":0.04442756387108196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3204758977","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5358698,0.00993621,0.39128563,0.03464859,0.00037052523,0.00028472798,0.0033186057,0.0024168435,0.021869114],"genre_scores_gemma":[0.9079836,0.000873866,0.087951705,0.0006618101,0.00019477325,0.00007505361,0.0014522118,0.00023165261,0.0005753721],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9813763,0.01041163,0.0011589131,0.0016564569,0.004738778,0.0006579998],"domain_scores_gemma":[0.8573274,0.10995721,0.008624624,0.01334102,0.0091844555,0.0015653298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026894566,0.001114984,0.0013592668,0.0031416828,0.0008337394,0.005486391,0.0013088052,0.0017047509,0.0011264859],"category_scores_gemma":[0.11361515,0.00038605902,0.0007819329,0.003662532,0.0017397435,0.009654511,0.0021278313,0.002712851,0.0010113539],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001640952,0.0004674724,0.17097235,0.0008745561,0.00041081157,0.00026254452,0.0011168262,0.054164153,0.010705863,0.028819455,0.012777231,0.7177878],"study_design_scores_gemma":[0.00018486984,0.0010332499,0.08195507,0.001070069,0.0004426359,0.001013252,0.003435455,0.6394223,0.04571796,0.18842559,0.03707409,0.00022554492],"about_ca_topic_score_codex":0.0021834383,"about_ca_topic_score_gemma":0.0034144141,"teacher_disagreement_score":0.026894566,"about_ca_system_score_codex":0.0016816565,"about_ca_system_score_gemma":0.0018129979,"threshold_uncertainty_score":0.14223379},"labels":[],"label_agreement":null},{"id":"W3204827391","doi":"10.5281/zenodo.5532640","title":"Scalable and Accurate Test Case Prioritization in Continuous Integration Contexts","year":2021,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; University of Ottawa","funders":"","keywords":"Prioritization; Scalability; Test (biology); Computer science; Data science; Database; Business; Process management; Geology","score_opus":0.021480443720523525,"score_gpt":0.2474873130833507,"score_spread":0.22600686936282718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3204827391","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26380563,0.0066894013,0.061255265,0.0014303627,0.0009223422,0.0010474137,0.57964265,0.06747122,0.01773581],"genre_scores_gemma":[0.15749103,0.00069177424,0.051443115,0.00034367046,0.00020581043,0.0007218934,0.78418696,0.0023838494,0.0025318314],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9921332,0.0018766973,0.0008023386,0.00252125,0.0020846794,0.00058191264],"domain_scores_gemma":[0.97822034,0.008305702,0.0015209907,0.006813832,0.0037978305,0.0013413654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050479784,0.002986487,0.0015412215,0.007201939,0.0009455632,0.00343505,0.0033059805,0.0018922304,0.0061425935],"category_scores_gemma":[0.02838391,0.00074333197,0.0015174219,0.0060651116,0.0007731874,0.0035333182,0.003641664,0.0019896652,0.006581608],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019211821,0.0020414467,0.0743497,0.003900413,0.0006356286,0.0008441077,0.0008265541,0.035208654,0.010278832,0.0055381777,0.6255239,0.23893133],"study_design_scores_gemma":[0.0013748861,0.0013117827,0.12687655,0.0014321512,0.00060751714,0.002098017,0.0019518014,0.2935575,0.028475313,0.033334352,0.5086505,0.00032962003],"about_ca_topic_score_codex":0.0064125014,"about_ca_topic_score_gemma":0.0127888145,"teacher_disagreement_score":0.007201939,"about_ca_system_score_codex":0.0012544475,"about_ca_system_score_gemma":0.0026244675,"threshold_uncertainty_score":0.026696563},"labels":[],"label_agreement":null},{"id":"W3206397610","doi":"10.1007/s10664-021-10083-5","title":"A fine-grained data set and analysis of tangling in bug fixing commits","year":2022,"lang":"en","type":"preprint","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University; University of Saskatchewan; University of British Columbia, Okanagan Campus; Kelowna General Hospital; University of British Columbia; University of Ottawa","funders":"Horizon 2020 Framework Programme; Technische Universität Clausthal; Deutsche Forschungsgemeinschaft","keywords":"Context (archaeology); Computer science; Software bug; Set (abstract data type); Code (set theory); Software; Source lines of code; Data mining; Programming language; Biology","score_opus":0.07673165410839632,"score_gpt":0.34358802997785276,"score_spread":0.26685637586945643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206397610","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97008944,0.00052561867,0.011073269,0.00036167077,0.000048331236,0.0002364,0.015461462,0.00037038172,0.0018334168],"genre_scores_gemma":[0.96159846,0.00009102571,0.013655967,0.00009147257,0.000037888716,0.00037433644,0.023522386,0.00011425798,0.0005141945],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.979675,0.0075764,0.002795614,0.0040859007,0.005221998,0.0006451155],"domain_scores_gemma":[0.68688923,0.19615611,0.045545705,0.03565111,0.03193932,0.0038185245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018041318,0.00046662518,0.00065450167,0.01178225,0.0011125803,0.001573871,0.0009788362,0.0014889835,0.0014291458],"category_scores_gemma":[0.13322534,0.00038103666,0.0004849198,0.008459627,0.001358744,0.0019275895,0.0024554983,0.0012980315,0.0007967242],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048055997,0.00041443214,0.9098739,0.000916429,0.00025864824,0.0007150537,0.0071425973,0.0064129042,0.0039795563,0.0017985372,0.011520951,0.05648641],"study_design_scores_gemma":[0.000039527506,0.00021282928,0.96428853,0.00027806446,0.000067518,0.00046296392,0.0023494896,0.015844401,0.0025029955,0.002293041,0.011576168,0.000084614534],"about_ca_topic_score_codex":0.0064612944,"about_ca_topic_score_gemma":0.007807063,"teacher_disagreement_score":0.018041318,"about_ca_system_score_codex":0.0010879164,"about_ca_system_score_gemma":0.0010803821,"threshold_uncertainty_score":0.09541273},"labels":[],"label_agreement":null},{"id":"W3206419348","doi":"10.1145/3485537","title":"Automatic migration from synchronous to asynchronous JavaScript APIs","year":2021,"lang":"en","type":"article","venue":"Proceedings of the ACM on Programming Languages","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Computer science; Asynchronous communication; JavaScript; Code refactoring; Callback; Programming language; Programmer; Distributed computing; Software; Computer network","score_opus":0.011510722755067895,"score_gpt":0.25955521929488073,"score_spread":0.24804449653981284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206419348","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3597583,0.00040732868,0.54971766,0.0007012903,0.00030978568,0.000669589,0.00074805226,0.08307887,0.004609177],"genre_scores_gemma":[0.55748004,0.0003687375,0.42224875,0.0006131803,0.00007670051,0.00034373964,0.0021087495,0.010927939,0.005832239],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9950594,0.0013069214,0.00061185396,0.0008097987,0.0017981828,0.00041387163],"domain_scores_gemma":[0.97203547,0.009218945,0.0019079285,0.011549327,0.004920026,0.00036826858],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004172956,0.0014514523,0.00058925763,0.0014759676,0.0007876825,0.0017037783,0.0030244773,0.0010672652,0.0012879156],"category_scores_gemma":[0.031453792,0.0011597505,0.0013215484,0.0011618829,0.0007861325,0.002747366,0.0021689455,0.0020604695,0.0013446708],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086716784,0.000810778,0.047125272,0.001041289,0.0001853237,0.00270766,0.006469895,0.023692869,0.18157394,0.01046037,0.019567853,0.7054976],"study_design_scores_gemma":[0.00028310795,0.00072618434,0.022569312,0.00039932577,0.00062432094,0.0024266692,0.0015047276,0.45505196,0.38461873,0.017187634,0.11431149,0.00029648782],"about_ca_topic_score_codex":0.0035028446,"about_ca_topic_score_gemma":0.0042410214,"teacher_disagreement_score":0.004172956,"about_ca_system_score_codex":0.00084792526,"about_ca_system_score_gemma":0.0026057952,"threshold_uncertainty_score":0.022068918},"labels":[],"label_agreement":null},{"id":"W3206481103","doi":"10.1007/s10664-021-10083-5","title":"A fine-grained data set and analysis of tangling in bug fixing commits","year":2022,"lang":"en","type":"article","venue":"Florence Research (University of Florence)","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University; University of Saskatchewan; University of British Columbia, Okanagan Campus; Kelowna General Hospital; University of British Columbia; University of Ottawa","funders":"Technische Universität Clausthal","keywords":"Context (archaeology); Computer science; Software bug; Code (set theory); Source lines of code; Set (abstract data type); Software; Data science; Data mining; Programming language; Biology","score_opus":0.10845759653231943,"score_gpt":0.32623385004434213,"score_spread":0.2177762535120227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206481103","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8971165,0.0010028603,0.019850545,0.0008912633,0.0001405002,0.00078305247,0.07444846,0.0013417731,0.004425074],"genre_scores_gemma":[0.8487356,0.00022797755,0.037891123,0.00025777082,0.00009830603,0.0014463757,0.10935813,0.00035170736,0.0016328504],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.982995,0.004630338,0.0024921147,0.003727503,0.005466074,0.00068897975],"domain_scores_gemma":[0.79451823,0.10909763,0.03199963,0.028125957,0.032359574,0.0038989743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012770875,0.00060535927,0.0006386513,0.012506995,0.0013560117,0.0016610294,0.0010535517,0.0017744495,0.0017916952],"category_scores_gemma":[0.100854106,0.0004129156,0.0006436608,0.008442476,0.001441308,0.0020925198,0.0027445129,0.0017295305,0.0013594417],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000774263,0.0007017996,0.83566445,0.0022128292,0.00033755554,0.0013560076,0.010466016,0.007493955,0.009024159,0.0028305782,0.03760218,0.09153633],"study_design_scores_gemma":[0.00007661821,0.00039138488,0.93889844,0.0005075619,0.00010237994,0.0009414486,0.003171223,0.0140162725,0.0049250447,0.0030594242,0.03376837,0.00014189615],"about_ca_topic_score_codex":0.007859104,"about_ca_topic_score_gemma":0.010990544,"teacher_disagreement_score":0.012770875,"about_ca_system_score_codex":0.0011285428,"about_ca_system_score_gemma":0.0016123494,"threshold_uncertainty_score":0.06753963},"labels":[],"label_agreement":null},{"id":"W3207869070","doi":"10.1016/j.jss.2021.111115","title":"Clone detection through srcClone: A program slicing based approach","year":2021,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Computer science; Program comprehension; clone (Java method); Benchmark (surveying); Program slicing; Scalability; Semantics (computer science); Source code; Software maintenance; Kernel (algebra); Precision and recall; Software; Programming language; Data mining; Code (set theory); Set (abstract data type); Software system; Artificial intelligence; Operating system; Biology","score_opus":0.022017679272204375,"score_gpt":0.26683964405038046,"score_spread":0.24482196477817608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207869070","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03294863,0.00030814682,0.92237025,0.000110828376,0.00004718262,0.00017877524,0.00051855395,0.04172433,0.0017932679],"genre_scores_gemma":[0.23116179,0.0001773597,0.7578131,0.00025407303,0.000054930795,0.000164319,0.0020530722,0.0046094493,0.003712011],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9974355,0.00032791935,0.00021612905,0.00057841005,0.0011964361,0.00024563703],"domain_scores_gemma":[0.99308467,0.002238359,0.0009587427,0.002064133,0.001492211,0.00016179262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014444175,0.0012570723,0.0015380709,0.0024869,0.0006468555,0.0014079337,0.001985749,0.0012270316,0.0047067637],"category_scores_gemma":[0.0047406983,0.00072162715,0.0017363284,0.0020435436,0.0008298892,0.0022549992,0.0016067964,0.001123497,0.00160349],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006700805,0.0003124168,0.016834034,0.0006393843,0.000348862,0.00088599,0.0006314491,0.021519126,0.2106437,0.009037236,0.010535108,0.72794265],"study_design_scores_gemma":[0.00010261725,0.00050198246,0.008365477,0.000110342604,0.00038993583,0.0013472218,0.00028224065,0.7332795,0.22906254,0.012570226,0.0138567835,0.00013112646],"about_ca_topic_score_codex":0.0034712828,"about_ca_topic_score_gemma":0.0066288244,"teacher_disagreement_score":0.0047067637,"about_ca_system_score_codex":0.00064739946,"about_ca_system_score_gemma":0.0019045525,"threshold_uncertainty_score":0.0157457},"labels":[],"label_agreement":null},{"id":"W3208243843","doi":"10.1109/rew53955.2021.00025","title":"Is BERT the New Silver Bullet? - An Empirical Investigation of Requirements Dependency Classification","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Transformer; Artificial intelligence; Silver bullet; Dependency (UML); Empirical research; Encoder; Machine learning; Data mining; Engineering","score_opus":0.10386123021689922,"score_gpt":0.35122460525203775,"score_spread":0.24736337503513853,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208243843","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9710811,0.0005160649,0.023544112,0.0007892976,0.00002785512,0.000120313736,0.0012985542,0.00028605637,0.0023366753],"genre_scores_gemma":[0.9813757,0.000101921825,0.014962364,0.00015028947,0.000019558709,0.00009058823,0.0027764803,0.000079674326,0.0004434652],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98990697,0.006138241,0.0007394235,0.0011755635,0.0017419884,0.00029787957],"domain_scores_gemma":[0.77637714,0.19744647,0.0074371374,0.012669361,0.0051648906,0.0009050024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0179953,0.00054422335,0.0004679122,0.0018487622,0.0006685978,0.0014173136,0.001177473,0.0010653894,0.002105535],"category_scores_gemma":[0.11619311,0.00031723044,0.0006594738,0.0022519554,0.0013464816,0.004820677,0.0013912718,0.0029734292,0.00077726674],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022481473,0.0027257237,0.51232255,0.0010368506,0.00039575412,0.00069339375,0.004092393,0.0654376,0.010464119,0.013185156,0.019290442,0.36810794],"study_design_scores_gemma":[0.00023193337,0.0013482219,0.2561051,0.00039770443,0.0001678383,0.001633391,0.0047375043,0.6671355,0.017332371,0.028322997,0.02243579,0.00015164184],"about_ca_topic_score_codex":0.003115072,"about_ca_topic_score_gemma":0.005857886,"teacher_disagreement_score":0.0179953,"about_ca_system_score_codex":0.0010772876,"about_ca_system_score_gemma":0.000615843,"threshold_uncertainty_score":0.095169365},"labels":[],"label_agreement":null},{"id":"W3208323417","doi":"10.1007/s10664-021-10045-x","title":"How do i refactor this? An empirical study on refactoring trends and topics in Stack Overflow","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Software engineering; Unit testing; Empirical research; Implementation; Software; Programming language","score_opus":0.04809255609342547,"score_gpt":0.3290085982567614,"score_spread":0.2809160421633359,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208323417","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9983246,0.00022549568,0.00045690036,0.00021135576,0.0000069520843,0.00002064588,0.00009046943,0.000030928888,0.00063272624],"genre_scores_gemma":[0.99664706,0.00043674835,0.0017621474,0.00011728238,0.000021870512,0.000025828072,0.00036267424,0.000036981517,0.0005893994],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99494874,0.0013970932,0.00070186035,0.00060402445,0.0019773003,0.0003709095],"domain_scores_gemma":[0.7850826,0.12681244,0.05749231,0.006175078,0.01851883,0.005918774],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.007555335,0.00033830566,0.0003359011,0.0041468325,0.0010443691,0.0021359874,0.0010238556,0.0011564542,0.0008664238],"category_scores_gemma":[0.121244974,0.00047086208,0.00043438136,0.003695196,0.0009257509,0.004660071,0.001093143,0.0021360088,0.0003090821],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020484076,0.00056143495,0.91184765,0.00018524042,0.00007731711,0.0003566733,0.028118916,0.000255124,0.0027551565,0.0004112467,0.00072659756,0.054499906],"study_design_scores_gemma":[0.000019049357,0.00050432136,0.9635402,0.00021513175,0.00012436356,0.00065992924,0.027272455,0.0021178343,0.002214794,0.0005993918,0.0026657234,0.000066694105],"about_ca_topic_score_codex":0.0062571196,"about_ca_topic_score_gemma":0.008850445,"teacher_disagreement_score":0.9924447,"about_ca_system_score_codex":0.0013392511,"about_ca_system_score_gemma":0.0024550776,"threshold_uncertainty_score":0.039956927},"labels":[],"label_agreement":null},{"id":"W3208730084","doi":"10.48550/arxiv.2110.14081","title":"A Controlled Experiment of Different Code Representations for Learning-Based Bug Repair","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Code (set theory); Source code; Representation (politics); Perspective (graphical); Point (geometry); Process (computing); Abstract syntax; Programming language; Syntax; Artificial intelligence; Set (abstract data type); Mathematics","score_opus":0.07040668802292974,"score_gpt":0.23906511192958935,"score_spread":0.1686584239066596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208730084","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9731006,0.0004394085,0.020018518,0.00028223218,0.00023873101,0.0012101638,0.001104172,0.0019422255,0.0016639461],"genre_scores_gemma":[0.94613564,0.00028737372,0.04503523,0.00044902938,0.00006805542,0.002519342,0.0020422197,0.00030871065,0.0031544408],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99764997,0.0008666157,0.00024002363,0.0007661683,0.00028686493,0.0001903969],"domain_scores_gemma":[0.96668065,0.02612286,0.0014884024,0.0033519056,0.0015205938,0.00083569554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00384535,0.0015135724,0.00065514137,0.000518261,0.00034543232,0.0008082222,0.0018308363,0.0019959023,0.00387854],"category_scores_gemma":[0.027241675,0.00052299246,0.00080969924,0.00038781835,0.0010558411,0.0022083563,0.0011853983,0.0028959573,0.0010459227],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.037432242,0.04696339,0.044887044,0.0059138737,0.0010559107,0.0013235608,0.004898876,0.28156796,0.13025802,0.008107062,0.022693308,0.41489878],"study_design_scores_gemma":[0.009237399,0.062286332,0.030450853,0.00075131946,0.001043849,0.0006383788,0.0015954806,0.72284466,0.13121028,0.014152637,0.025393084,0.00039581553],"about_ca_topic_score_codex":0.0017068427,"about_ca_topic_score_gemma":0.002255333,"teacher_disagreement_score":0.00387854,"about_ca_system_score_codex":0.0006143589,"about_ca_system_score_gemma":0.0009710259,"threshold_uncertainty_score":0.02033639},"labels":[],"label_agreement":null},{"id":"W3208913301","doi":"10.32920/ryerson.14638470.v1","title":"Partially observable Markov Decision Process to prioritize software defects","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Partially observable Markov decision process; Dependency (UML); Dependency graph; Exploit; Software bug; Prioritization; Software; Graph; Process (computing); Software quality; Software regression; Data mining; Risk analysis (engineering); Markov chain; Markov model; Software development; Machine learning; Artificial intelligence; Computer security; Engineering; Theoretical computer science","score_opus":0.025687265362385556,"score_gpt":0.2996946730810493,"score_spread":0.2740074077186637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208913301","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0643449,0.00037279818,0.9299793,0.0007130605,0.0000981587,0.000280298,0.00047093353,0.00054997805,0.0031904853],"genre_scores_gemma":[0.90397125,0.0003040143,0.09211806,0.00021079704,0.00005233215,0.00047528293,0.0004958757,0.000045795387,0.0023266047],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99723166,0.0010189217,0.00014005158,0.0006196949,0.00056309334,0.00042651143],"domain_scores_gemma":[0.98432195,0.012832988,0.0011563625,0.00022832277,0.0010211553,0.00043917203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032714957,0.0016301847,0.001869381,0.0012275951,0.0006799754,0.0014091147,0.0014940625,0.0014450949,0.0037938212],"category_scores_gemma":[0.010737421,0.0008337819,0.0015206502,0.0009363327,0.0014096837,0.001389338,0.0012169895,0.0024517856,0.00028911428],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016361738,0.00006622217,0.002365075,0.00013392106,0.00006179725,0.00017434313,0.000087197055,0.97409874,0.0005330836,0.01298065,0.0004959079,0.008839363],"study_design_scores_gemma":[0.00002480386,0.000036371486,0.00024555004,0.0000096395615,0.000019954596,0.000011596611,0.000011662201,0.9936441,0.00015738598,0.005691555,0.00013954587,0.000007800476],"about_ca_topic_score_codex":0.02224345,"about_ca_topic_score_gemma":0.01455428,"teacher_disagreement_score":0.02224345,"about_ca_system_score_codex":0.0022320237,"about_ca_system_score_gemma":0.0035682977,"threshold_uncertainty_score":0.044227958},"labels":[],"label_agreement":null},{"id":"W3209001851","doi":"10.1016/j.infsof.2021.106756","title":"Early Prediction for Merged vs Abandoned Code Changes in Modern Code Reviews","year":2019,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Code review; Code (set theory); Machine learning; Source code; Artificial intelligence; Classifier (UML); Suite; Software; Source lines of code; Empirical research; Software inspection; Process (computing); Software engineering; Software development; Software quality; Programming language; Statistics","score_opus":0.07216682349375034,"score_gpt":0.20648882950485375,"score_spread":0.13432200601110342,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209001851","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.973486,0.0020351966,0.019257305,0.00048220836,0.00011874202,0.00016815226,0.0016420376,0.0010720611,0.0017382818],"genre_scores_gemma":[0.98338056,0.00038845773,0.012744841,0.00008287837,0.000070760216,0.00006179937,0.0023178426,0.000064508764,0.00088841055],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9932462,0.0013654098,0.00077415584,0.0013818304,0.0027536927,0.00047871855],"domain_scores_gemma":[0.913725,0.044451095,0.021259364,0.0024176976,0.01608178,0.002065039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047012563,0.0007385587,0.00061242824,0.0065319245,0.00056523795,0.0020568958,0.00088944327,0.0011640054,0.0006679622],"category_scores_gemma":[0.04679914,0.00033398013,0.00046322952,0.002068169,0.0004540118,0.0022181156,0.00089358055,0.0012347496,0.0008869721],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039298323,0.00021292748,0.8689314,0.00034941253,0.00010346266,0.0005229089,0.00065953325,0.0062039043,0.0037814276,0.00026491188,0.0051220837,0.113455124],"study_design_scores_gemma":[0.00003937469,0.00071524386,0.72732204,0.00030028893,0.00019366066,0.0014417842,0.0012969687,0.24330339,0.01626111,0.001132612,0.00788464,0.000108880326],"about_ca_topic_score_codex":0.003473512,"about_ca_topic_score_gemma":0.007838986,"teacher_disagreement_score":0.0065319245,"about_ca_system_score_codex":0.0007717105,"about_ca_system_score_gemma":0.00096977735,"threshold_uncertainty_score":0.024862945},"labels":[],"label_agreement":null},{"id":"W3209347229","doi":"10.1002/smr.2402","title":"A systematic mapping study on the employment of neural networks on software engineering projects: Where to go next?","year":2021,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Artificial neural network; Domain (mathematical analysis); Computer science; Deep learning; Software; Field (mathematics); Software engineering; Artificial intelligence; Software development; Social software engineering; Point (geometry); Data science; Software construction","score_opus":0.028202629790664114,"score_gpt":0.2700500468416283,"score_spread":0.2418474170509642,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209347229","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97859484,0.0018377204,0.015198428,0.0004832626,0.000045732417,0.00020007566,0.0006758375,0.00007084526,0.0028932253],"genre_scores_gemma":[0.9870253,0.0005956464,0.01070166,0.00009167828,0.000016132622,0.00021646889,0.00067959825,0.000026582222,0.0006468928],"study_design_codex":"observational","study_design_gemma":"systematic_review","domain_scores_codex":[0.9877451,0.007789139,0.00094658486,0.001411137,0.0017887576,0.00031927697],"domain_scores_gemma":[0.85581595,0.087147765,0.01638448,0.011280475,0.028113568,0.0012577248],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011135331,0.00042147186,0.00035972154,0.0067945193,0.00086729904,0.0014050083,0.0007653995,0.0006239227,0.0010696877],"category_scores_gemma":[0.07824511,0.00027334224,0.0005687103,0.006414637,0.0009883045,0.0027665142,0.0019415462,0.0008188542,0.00041087184],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003423904,0.0006184114,0.6168388,0.0017901281,0.00053587736,0.0004947869,0.015502799,0.0037956983,0.0036118086,0.007521585,0.0041030208,0.34484476],"study_design_scores_gemma":[0.000099826335,0.0018823188,0.80588436,0.0036659862,0.000985383,0.0009995517,0.051396586,0.05772591,0.01742451,0.023258246,0.036447085,0.00023026836],"about_ca_topic_score_codex":0.003402834,"about_ca_topic_score_gemma":0.006114685,"teacher_disagreement_score":0.98886466,"about_ca_system_score_codex":0.0010242703,"about_ca_system_score_gemma":0.0015157105,"threshold_uncertainty_score":0.058889925},"labels":[],"label_agreement":null},{"id":"W3209468076","doi":"10.1145/3530785","title":"On Wasted Contributions: Understanding the Dynamics of Contributor-Abandoned Pull Requests–A Mixed-Methods Study of 10 Large Open-Source Projects","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Concordia University","funders":"","keywords":"Computer science; Open source; Dynamics (music); Open source software; Management science; Operations research; Software; Economics; Sociology; Engineering; Programming language","score_opus":0.08502718441017452,"score_gpt":0.3720169809182515,"score_spread":0.286989796508077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209468076","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9905545,0.0008787806,0.0063423566,0.0003167501,0.000019223438,0.00023391978,0.00075276126,0.000056990964,0.0008446557],"genre_scores_gemma":[0.9873848,0.00041900962,0.009071165,0.00018074735,0.000044264783,0.0004910801,0.0013640112,0.000081883445,0.00096303027],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9656782,0.022567462,0.0025044759,0.0037934156,0.004324904,0.0011315872],"domain_scores_gemma":[0.55422914,0.36320633,0.045820754,0.012448997,0.020635838,0.003658929],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.047358554,0.00054650276,0.0007350277,0.006567546,0.0013439935,0.0031876748,0.001685757,0.0013223572,0.0009841243],"category_scores_gemma":[0.17870174,0.0006034715,0.00090782036,0.0045151543,0.0013215622,0.0047545996,0.0025691777,0.001416825,0.0005716469],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075856294,0.0007369439,0.86121756,0.0016792149,0.0004759785,0.0007572922,0.05649165,0.0013661883,0.0027715305,0.0019794176,0.0027182442,0.06904727],"study_design_scores_gemma":[0.000087056134,0.0010351517,0.8791635,0.0007802125,0.00030811937,0.0011290279,0.059321575,0.038918436,0.002858546,0.0029345548,0.013247912,0.0002158533],"about_ca_topic_score_codex":0.0047682384,"about_ca_topic_score_gemma":0.0084212795,"teacher_disagreement_score":0.9526414,"about_ca_system_score_codex":0.0016120571,"about_ca_system_score_gemma":0.0019515789,"threshold_uncertainty_score":0.25045896},"labels":[],"label_agreement":null},{"id":"W3209827742","doi":"10.48550/arxiv.2102.09775","title":"Characterizing and Mitigating Self-Admitted Technical Debt in Build Systems","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Technical debt; Computer science; Artifact (error); Software; Set (abstract data type); Software system; Data science; Dependency (UML); Data mining; Software development; Software engineering; Artificial intelligence","score_opus":0.03750002698376278,"score_gpt":0.1958700061272583,"score_spread":0.15836997914349552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209827742","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9740538,0.00037546078,0.023342995,0.00027578228,0.000015363163,0.000112338814,0.00025824638,0.00042487096,0.0011412019],"genre_scores_gemma":[0.9780961,0.00018172649,0.01993878,0.00008771024,0.000020384054,0.000079835954,0.0007737907,0.0001160934,0.0007055412],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9893884,0.0033196693,0.001302088,0.0012506054,0.004108869,0.0006303144],"domain_scores_gemma":[0.9010954,0.0500039,0.028331231,0.0055760164,0.013517802,0.001475642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072623687,0.00061138783,0.00046121926,0.0064283097,0.0012875833,0.0026868887,0.00080317,0.0010607222,0.000420143],"category_scores_gemma":[0.06771243,0.00053802686,0.00044164425,0.0029107777,0.0011509319,0.0039140745,0.0026251178,0.0010696214,0.00026701423],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022055453,0.00019022527,0.79354376,0.00089960074,0.00009067817,0.001206391,0.0335534,0.0041952166,0.023404349,0.0014098894,0.0016681117,0.13961786],"study_design_scores_gemma":[0.000020411959,0.00043335784,0.85160804,0.00057285995,0.00013628097,0.002154539,0.023143057,0.08551773,0.017686868,0.004312295,0.014236222,0.0001783693],"about_ca_topic_score_codex":0.0054235677,"about_ca_topic_score_gemma":0.010436288,"teacher_disagreement_score":0.0072623687,"about_ca_system_score_codex":0.0013080861,"about_ca_system_score_gemma":0.001499343,"threshold_uncertainty_score":0.038407505},"labels":[],"label_agreement":null},{"id":"W3210111894","doi":"10.1109/seaa53835.2021.00015","title":"Automated Support for Searching and Selecting Evidence in Software Engineering: A Cross-domain Systematic Mapping","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Workload; Computer science; Automation; Domain (mathematical analysis); Context (archaeology); Field (mathematics); Data science; Systematic review; Variety (cybernetics); Selection (genetic algorithm); Software engineering; Software; Precision and recall; Information retrieval; Machine learning; Artificial intelligence; Engineering; MEDLINE","score_opus":0.042381102224931404,"score_gpt":0.32087167509287495,"score_spread":0.27849057286794354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3210111894","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13151015,0.3441812,0.416418,0.017284758,0.001096716,0.06434633,0.009616375,0.0019607877,0.013585673],"genre_scores_gemma":[0.22309107,0.061616264,0.6759015,0.0023730465,0.00018975112,0.03260939,0.0034076825,0.00019379501,0.00061749836],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.71500057,0.17896594,0.069058985,0.009288463,0.02644474,0.0012414025],"domain_scores_gemma":[0.24654931,0.6271023,0.040127467,0.038255837,0.046142265,0.0018227781],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.24634156,0.0018984239,0.0040564216,0.067669496,0.0033278563,0.0076831416,0.003569822,0.002545024,0.0051013185],"category_scores_gemma":[0.5345103,0.0016624145,0.0063445256,0.031889107,0.0028292092,0.012162623,0.01242479,0.0023700937,0.000891044],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004923722,0.000330259,0.012788538,0.21487321,0.005286848,0.00042690174,0.015108111,0.0015634296,0.0022096264,0.0068309684,0.00376673,0.7363231],"study_design_scores_gemma":[0.0015113942,0.002328778,0.038962368,0.700425,0.031740993,0.0021220979,0.022544453,0.010192925,0.009557545,0.059980623,0.11999299,0.00064083334],"about_ca_topic_score_codex":0.0033062634,"about_ca_topic_score_gemma":0.009012227,"teacher_disagreement_score":0.7536584,"about_ca_system_score_codex":0.006468514,"about_ca_system_score_gemma":0.035769396,"threshold_uncertainty_score":0.929395},"labels":[],"label_agreement":null},{"id":"W3210853439","doi":"10.5281/zenodo.5361067","title":"How Do I Refactor This? An Empirical Study on Refactoring Trends and Topics in Stack Overflow","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Code refactoring; Stack (abstract data type); Computer science; Empirical research; Programming language; Mathematics; Software; Statistics","score_opus":0.062116709552840176,"score_gpt":0.3083726042637263,"score_spread":0.24625589471088613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3210853439","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99703217,0.00022002576,0.00057827594,0.00046883302,0.000010598375,0.00006957883,0.0005451242,0.000029738509,0.0010456949],"genre_scores_gemma":[0.9937651,0.00046646572,0.0025331448,0.00034425396,0.000025682113,0.00022975732,0.0014635447,0.000055823893,0.0011161916],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99497354,0.002063406,0.0006453108,0.00069796696,0.0011770789,0.0004426821],"domain_scores_gemma":[0.8695428,0.088057555,0.023253009,0.0033615385,0.01182504,0.003960133],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0108227,0.0002980434,0.00035025072,0.003371776,0.0014526608,0.0019338403,0.0009383087,0.0010594528,0.0021211093],"category_scores_gemma":[0.07551829,0.00039410027,0.00034560476,0.0029098075,0.0013297176,0.0045732,0.0017115609,0.0019793955,0.00055510097],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003003307,0.00091630017,0.6374008,0.0008763362,0.0000524126,0.0006617551,0.29579034,0.0003128795,0.0032999176,0.0010750886,0.00614711,0.053166747],"study_design_scores_gemma":[0.000035464247,0.00040130247,0.7153026,0.0006214578,0.00005537142,0.00034370128,0.2593821,0.0018389189,0.0018877009,0.0008461984,0.019145472,0.00013972551],"about_ca_topic_score_codex":0.008084727,"about_ca_topic_score_gemma":0.013365526,"teacher_disagreement_score":0.9966282,"about_ca_system_score_codex":0.0022002982,"about_ca_system_score_gemma":0.0019908368,"threshold_uncertainty_score":0.05723661},"labels":[],"label_agreement":null},{"id":"W3210938546","doi":"10.3390/app11136188","title":"Automated Extraction and Time-Cost Prediction of Contractual Reporting Requirements in Construction Using Natural Language Processing and Simulation","year":2021,"lang":"en","type":"article","venue":"Applied Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Workflow; Computer science; Overhead (engineering); Identification (biology); Database; Programming language","score_opus":0.037233583633287555,"score_gpt":0.3453840418439745,"score_spread":0.30815045821068693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3210938546","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19979756,0.00015379705,0.79163426,0.00039413446,0.000026744481,0.00040641104,0.0018461229,0.003484986,0.0022558914],"genre_scores_gemma":[0.46847242,0.00012893674,0.52693313,0.000048099115,0.000010236741,0.00045450963,0.0031091513,0.00015609236,0.00068748323],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970067,0.0013624533,0.00030085634,0.00037826222,0.0008318373,0.00011994203],"domain_scores_gemma":[0.98221946,0.0124374395,0.002292051,0.0010495717,0.0018326024,0.00016892234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037019204,0.00071453647,0.00047440725,0.0019761748,0.00048022333,0.0010723339,0.0008478622,0.0006014681,0.0011917566],"category_scores_gemma":[0.016591271,0.00045332545,0.00086516526,0.0013147183,0.0004073094,0.0012872486,0.0006683446,0.00084134936,0.00033343738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027804633,0.0003023223,0.02137607,0.00039383097,0.00006872016,0.00036039116,0.0006302374,0.77142686,0.010130504,0.012522519,0.00456297,0.17794754],"study_design_scores_gemma":[0.0000090796675,0.000023974904,0.0019033088,0.0000142676845,0.0000072508396,0.000032809778,0.00005946479,0.9917972,0.0027890727,0.0020748235,0.0012755851,0.000013134433],"about_ca_topic_score_codex":0.018149128,"about_ca_topic_score_gemma":0.020987032,"teacher_disagreement_score":0.018149128,"about_ca_system_score_codex":0.0017946698,"about_ca_system_score_gemma":0.0035395508,"threshold_uncertainty_score":0.036086977},"labels":[],"label_agreement":null},{"id":"W3211076835","doi":"10.1109/rew53955.2021.00012","title":"Generating Sequence Diagram from Natural Language Requirements","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Alberta Innovates","keywords":"Sequence diagram; Computer science; Unified Modeling Language; Communication diagram; Correctness; Completeness (order theory); Class diagram; Programming language; Applications of UML; Use Case Diagram; UML tool; Natural language; Sequence (biology); Natural language processing; Software engineering; Software","score_opus":0.03411350203033933,"score_gpt":0.3147973851656794,"score_spread":0.2806838831353401,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211076835","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02398839,0.00007162737,0.96492785,0.00016080817,0.00004702036,0.0011634122,0.0021939476,0.0046925796,0.002754335],"genre_scores_gemma":[0.07105061,0.00014427287,0.9172739,0.0000650192,0.000013979558,0.0012780734,0.008207398,0.0004999035,0.0014668067],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9959176,0.0016831016,0.00035698691,0.00050033693,0.0014429402,0.00009903106],"domain_scores_gemma":[0.98680484,0.008590843,0.00064057583,0.0011959557,0.0026438108,0.00012408801],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024335957,0.0012497444,0.0004904353,0.003326817,0.00042044595,0.00080742314,0.00082345394,0.0008431262,0.0046061166],"category_scores_gemma":[0.017119287,0.00040448204,0.0013459636,0.0010789699,0.00035949497,0.00092289614,0.00074699114,0.0006740517,0.001616603],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066975254,0.0006629168,0.006733688,0.0031174242,0.00014463707,0.003266871,0.0025550935,0.18218246,0.11049171,0.05647547,0.020350827,0.61334914],"study_design_scores_gemma":[0.00022785191,0.0004387073,0.0024916264,0.00031064585,0.0000935569,0.0014050562,0.0005856649,0.8084516,0.09065524,0.028805448,0.06642384,0.00011070378],"about_ca_topic_score_codex":0.0029317834,"about_ca_topic_score_gemma":0.0036296733,"teacher_disagreement_score":0.0046061166,"about_ca_system_score_codex":0.00076951843,"about_ca_system_score_gemma":0.0027124637,"threshold_uncertainty_score":0.015409052},"labels":[],"label_agreement":null},{"id":"W3211582933","doi":"10.54590/pop.2021.004","title":"Social Analytics Through Spyral","year":2021,"lang":"en","type":"article","venue":"Pop! Public Open Participatory","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"JavaScript; Computer science; Analytics; World Wide Web; Development environment; Code (set theory); Data science; Software engineering; Programming language; Human–computer interaction","score_opus":0.27338811428755994,"score_gpt":0.41857461050994715,"score_spread":0.14518649622238722,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211582933","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01965612,0.0010954738,0.8225886,0.006865758,0.0005066999,0.0005296304,0.004169431,0.027424973,0.11716335],"genre_scores_gemma":[0.32966074,0.0016349049,0.6054295,0.0014185698,0.00050048606,0.0013896498,0.008139658,0.0053775413,0.046448924],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9912894,0.005057769,0.00035639,0.0014320959,0.0015529029,0.000311471],"domain_scores_gemma":[0.9864818,0.007869219,0.0006594409,0.0036062582,0.00068970147,0.0006935286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007798849,0.0010508031,0.0005444968,0.0043244846,0.0026089933,0.009230997,0.0021226045,0.0011124114,0.029748157],"category_scores_gemma":[0.018323813,0.00066033076,0.0011815255,0.0028927552,0.004663002,0.019964695,0.017669698,0.0019121369,0.007790212],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032105384,0.0001532647,0.00330795,0.0009779858,0.00008060392,0.00064427475,0.027381213,0.003434669,0.003255125,0.5567602,0.05724561,0.34643802],"study_design_scores_gemma":[0.00004297321,0.000052633906,0.001187319,0.0003325409,0.00001630743,0.0002896809,0.005074392,0.01192857,0.002196993,0.26822,0.7106036,0.000054957753],"about_ca_topic_score_codex":0.0016869684,"about_ca_topic_score_gemma":0.0029264954,"teacher_disagreement_score":0.029748157,"about_ca_system_score_codex":0.0013344848,"about_ca_system_score_gemma":0.0016820568,"threshold_uncertainty_score":0.099517524},"labels":[],"label_agreement":null},{"id":"W3211768873","doi":"10.1145/3528579.3529177","title":"How developers and managers define and trade productivity for quality","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Productivity; Quality (philosophy); Software; Software quality; Business; Robustness (evolution); Computer science; Software development; Knowledge management; Economics","score_opus":0.06251090589772802,"score_gpt":0.31121229572994313,"score_spread":0.24870138983221513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211768873","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3314315,0.015002921,0.07405576,0.28350317,0.0010742148,0.000174162,0.00019058389,0.0005021287,0.2940656],"genre_scores_gemma":[0.9807997,0.0018087486,0.0062713274,0.0051666563,0.0001983943,0.00008169184,0.000038270526,0.0001767368,0.005458427],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9564668,0.029595176,0.0014466848,0.0029215603,0.006624331,0.0029454713],"domain_scores_gemma":[0.8510671,0.10683186,0.013824552,0.006541052,0.013865246,0.007870107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03318148,0.00047383737,0.0005905283,0.0029719113,0.0038752214,0.030903915,0.001996627,0.0056073642,0.0073922826],"category_scores_gemma":[0.12913151,0.00074256584,0.00037333585,0.0033258894,0.018123893,0.027952151,0.006929048,0.004577287,0.0015300448],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016756002,0.00016972785,0.033276968,0.000623136,0.00011983989,0.00043817097,0.20859101,0.0016830346,0.0013895805,0.6191687,0.031896215,0.102476],"study_design_scores_gemma":[0.0000844947,0.00011203232,0.0205323,0.0010211974,0.000070427224,0.00037986622,0.15152407,0.0026264854,0.0014689546,0.67927474,0.14277762,0.00012780506],"about_ca_topic_score_codex":0.0059907897,"about_ca_topic_score_gemma":0.0057358653,"teacher_disagreement_score":0.03318148,"about_ca_system_score_codex":0.00870465,"about_ca_system_score_gemma":0.011251762,"threshold_uncertainty_score":0.17548251},"labels":[],"label_agreement":null},{"id":"W3211925781","doi":"10.1145/3491211","title":"An Empirical Study of the Effectiveness of an Ensemble of Stand-alone Sentiment Detection Tools for Software Engineering Datasets","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Saskatchewan; Concordia University; University of Calgary","funders":"","keywords":"Computer science; Sentiment analysis; Software; Majority rule; Artificial intelligence; Machine learning; Empirical research; Detector; Data mining","score_opus":0.0814680668414065,"score_gpt":0.3610115127954521,"score_spread":0.2795434459540456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211925781","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.943196,0.0045340112,0.021558283,0.0007559565,0.0005173709,0.00090922817,0.022640996,0.0023525315,0.0035357587],"genre_scores_gemma":[0.7963675,0.0010067715,0.08142878,0.00039795498,0.00033118145,0.0008456191,0.117794424,0.00020964882,0.0016180889],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99256,0.0026872454,0.0010593253,0.0016960952,0.0015936858,0.00040361012],"domain_scores_gemma":[0.9788519,0.011313766,0.0013011462,0.002693726,0.0050651133,0.00077436096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015920976,0.0019057917,0.0015332126,0.004577098,0.0012011977,0.0017054189,0.0015052274,0.00157105,0.0007482875],"category_scores_gemma":[0.02032055,0.00032974067,0.001771274,0.003305621,0.0007376197,0.0028864173,0.0017134781,0.0016744633,0.0009227493],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035145415,0.003937093,0.43388045,0.0028286255,0.004152777,0.0005851057,0.0012483855,0.03728474,0.025140656,0.0012219655,0.06610791,0.42009774],"study_design_scores_gemma":[0.0007220311,0.0045704283,0.36932158,0.0005284521,0.0019228252,0.0022737978,0.0029115425,0.5461591,0.036497828,0.002833791,0.031965096,0.00029355142],"about_ca_topic_score_codex":0.0029433907,"about_ca_topic_score_gemma":0.0050652656,"teacher_disagreement_score":0.015920976,"about_ca_system_score_codex":0.0009285034,"about_ca_system_score_gemma":0.00082476344,"threshold_uncertainty_score":0.08419919},"labels":[],"label_agreement":null},{"id":"W3212549866","doi":"10.18653/v1/2021.findings-emnlp.214","title":"Novel Natural Language Summarization of Program Code via Leveraging Multiple Input Representations","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Automatic summarization; Programming language; Python (programming language); Source code; Code (set theory); Encoder; Redundant code; Code generation; Task (project management); Artificial intelligence; Natural language processing; Set (abstract data type); Operating system","score_opus":0.024817799380612772,"score_gpt":0.31090854281459185,"score_spread":0.28609074343397906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3212549866","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060064297,0.0009238039,0.9072844,0.0006105231,0.0001441495,0.00023766235,0.0025663152,0.026145814,0.0020229388],"genre_scores_gemma":[0.30472195,0.00049894647,0.66612995,0.00049553247,0.00014944217,0.0004029461,0.019606935,0.0014171108,0.006577106],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992078,0.0001988352,0.000048095637,0.00031788566,0.00015782651,0.000069566864],"domain_scores_gemma":[0.9983223,0.0006876244,0.00021763268,0.00029140475,0.00040559782,0.00007538231],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007138365,0.001660027,0.00064976234,0.0016518302,0.00034991582,0.00078682386,0.0013898701,0.0010232326,0.0020514326],"category_scores_gemma":[0.00419336,0.00035620527,0.0012009032,0.0010287322,0.00043583222,0.0028406072,0.0011404869,0.0018858031,0.0014348214],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004140016,0.00041769698,0.00295788,0.0007003314,0.00013754336,0.00045534098,0.0005247299,0.0659845,0.064394236,0.0052009188,0.023720592,0.8350922],"study_design_scores_gemma":[0.000037178448,0.00026589265,0.0015885998,0.000043291595,0.000078090256,0.00022195937,0.00017020504,0.94206166,0.03523741,0.009923287,0.010335368,0.000037160436],"about_ca_topic_score_codex":0.0037829552,"about_ca_topic_score_gemma":0.0071153287,"teacher_disagreement_score":0.0037829552,"about_ca_system_score_codex":0.0007287251,"about_ca_system_score_gemma":0.0012427617,"threshold_uncertainty_score":0.0075218678},"labels":[],"label_agreement":null},{"id":"W3213599224","doi":"10.1109/vissoft52517.2021.00027","title":"Understanding High-Level Behavior with a Light-Traces Visualization Metaphor","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Animation; Computer science; Visualization; Metaphor; Exploit; Set (abstract data type); Human–computer interaction; Information visualization; Computer animation; Data visualization; Computer graphics (images); Artificial intelligence; Programming language","score_opus":0.11658962961440078,"score_gpt":0.3044504734519165,"score_spread":0.18786084383751572,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3213599224","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018147467,0.000115840005,0.9736009,0.00063702045,0.000029750854,0.00006168548,0.00010371667,0.0038591763,0.00344438],"genre_scores_gemma":[0.30228746,0.00044925863,0.69190055,0.0002187477,0.000029883286,0.00016786737,0.00022053358,0.0010224129,0.0037032836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996314,0.00015462509,0.000020683201,0.000059745264,0.000090933296,0.00004254425],"domain_scores_gemma":[0.9981223,0.00115296,0.00013808951,0.00024239857,0.00017293007,0.0001713191],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009946873,0.001252962,0.00034610112,0.0014432404,0.00043150788,0.0024383038,0.0009847381,0.0011149712,0.0060894703],"category_scores_gemma":[0.003931342,0.00049725,0.000684039,0.00045169066,0.0012468683,0.0040629655,0.0016014812,0.0014857351,0.00095004815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010780464,0.00052569254,0.007381233,0.0019458911,0.00014844748,0.0030190824,0.03190985,0.057677906,0.29377544,0.28330985,0.01842591,0.30080265],"study_design_scores_gemma":[0.00024064168,0.0005935316,0.004472376,0.0007236562,0.00016035755,0.0029097088,0.0050105816,0.4582446,0.11824525,0.19676495,0.21233632,0.00029807066],"about_ca_topic_score_codex":0.0019380556,"about_ca_topic_score_gemma":0.0019137876,"teacher_disagreement_score":0.0060894703,"about_ca_system_score_codex":0.00039834864,"about_ca_system_score_gemma":0.0006032734,"threshold_uncertainty_score":0.020371318},"labels":[],"label_agreement":null},{"id":"W3214151424","doi":"10.48550/arxiv.2111.07263","title":"Code Representation Learning with Prüfer Sequences","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Automatic summarization; Sequence (biology); Representation (politics); Source code; Artificial intelligence; Encoding (memory); Syntax; Benchmark (surveying); Deep learning; Code (set theory); Program comprehension; Natural language processing; Abstract syntax; Programming language; Software; Software system","score_opus":0.08707538642658927,"score_gpt":0.21679399147137887,"score_spread":0.1297186050447896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214151424","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07402587,0.0005497907,0.9145074,0.00063345255,0.00007917163,0.00010221815,0.001115122,0.007007615,0.0019793408],"genre_scores_gemma":[0.62579626,0.0005300131,0.35832617,0.0003637081,0.000082287115,0.00039096802,0.006833058,0.00047134343,0.0072061643],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99939287,0.0001597191,0.00003964939,0.00019989464,0.0001453104,0.00006251273],"domain_scores_gemma":[0.99827254,0.000602194,0.00021449488,0.00046205515,0.00038199977,0.000066751396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006616656,0.0009165861,0.0005661454,0.0013585985,0.0003327397,0.0007893562,0.0014373165,0.0010896454,0.002269393],"category_scores_gemma":[0.0052917902,0.0003006711,0.0006531429,0.0013078039,0.0006645414,0.0032885785,0.0010011523,0.0019302095,0.001107998],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002811675,0.00018820378,0.0019280486,0.00019607693,0.000043941203,0.00014648061,0.0001671555,0.41713983,0.007746907,0.025287662,0.009621258,0.5372532],"study_design_scores_gemma":[0.000014906116,0.00006770456,0.00014968975,0.000013272346,0.000008199,0.00003050825,0.000017754452,0.9723672,0.004286751,0.020905314,0.002129621,0.000009086075],"about_ca_topic_score_codex":0.004890773,"about_ca_topic_score_gemma":0.006453998,"teacher_disagreement_score":0.004890773,"about_ca_system_score_codex":0.0011568305,"about_ca_system_score_gemma":0.0015506713,"threshold_uncertainty_score":0.009724617},"labels":[],"label_agreement":null},{"id":"W3214476327","doi":"10.3390/sym13112166","title":"Software Defect Prediction Using Wrapper Feature Selection Based on Dynamic Re-Ranking Strategy","year":2021,"lang":"en","type":"article","venue":"Symmetry","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Yayasan UTP; Universiti Teknologi Petronas","keywords":"Computer science; Feature selection; Overfitting; Data mining; Maxima and minima; Curse of dimensionality; Machine learning; Ranking (information retrieval); Artificial intelligence; Software; Classifier (UML); Feature (linguistics); Process (computing); Pattern recognition (psychology); Artificial neural network; Mathematics","score_opus":0.01617267138903952,"score_gpt":0.27165605243976,"score_spread":0.2554833810507205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214476327","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2102952,0.0007840623,0.7831616,0.0002078789,0.000096368814,0.00020127342,0.00040552014,0.0032421553,0.001606007],"genre_scores_gemma":[0.89886284,0.00024508158,0.09746793,0.00008863486,0.000055811914,0.00021371178,0.0009992551,0.00008579519,0.0019808095],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992182,0.0000880726,0.000062889965,0.00019199838,0.00032496385,0.00011383569],"domain_scores_gemma":[0.9987174,0.0003681181,0.00013588759,0.00011247249,0.000600391,0.00006572724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000725688,0.0012499731,0.0016332632,0.0028746608,0.00039954548,0.0007399709,0.001053895,0.00060753105,0.001026103],"category_scores_gemma":[0.0023499776,0.00032239317,0.0011763284,0.0013205851,0.00026703614,0.0008019334,0.00051028957,0.00043093422,0.00047018952],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003628037,0.0004077756,0.020810772,0.00013670318,0.00021970944,0.00062240474,0.00013828294,0.23552698,0.021118702,0.0010963944,0.0056709764,0.71388847],"study_design_scores_gemma":[0.000017647002,0.000118027914,0.003835875,0.000008630969,0.000047674355,0.00012209117,0.00002489464,0.99137306,0.0032976093,0.00054533547,0.00059057475,0.0000187152],"about_ca_topic_score_codex":0.007324806,"about_ca_topic_score_gemma":0.005101036,"teacher_disagreement_score":0.007324806,"about_ca_system_score_codex":0.00042432704,"about_ca_system_score_gemma":0.000903585,"threshold_uncertainty_score":0.014564335},"labels":[],"label_agreement":null},{"id":"W3214674962","doi":"10.1109/models50736.2021.00014","title":"Repository Mining for Changes in Simulink Models","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Artifact (error); Computer science; Automotive industry; Software; Software engineering; Model-based design; Software development; Simulation; Engineering; Artificial intelligence; Operating system","score_opus":0.04334033440436619,"score_gpt":0.28583879585364513,"score_spread":0.24249846144927895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214674962","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48536876,0.0034631826,0.3970088,0.002236908,0.00071032176,0.0017169415,0.051238686,0.045801856,0.012454551],"genre_scores_gemma":[0.54837453,0.0013844782,0.3289093,0.00022456027,0.00007732847,0.00072436256,0.11186583,0.0030671791,0.005372386],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98878074,0.0015384261,0.0018099349,0.0022735987,0.005115376,0.00048193015],"domain_scores_gemma":[0.9584879,0.014561047,0.006619967,0.010951746,0.0086358525,0.0007434927],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006046161,0.0015504586,0.001092484,0.0138086835,0.0012376639,0.0041404455,0.003853218,0.0014337492,0.0025812136],"category_scores_gemma":[0.053886767,0.0010585679,0.0024370225,0.008794612,0.00092300295,0.00484156,0.0034817278,0.0018538993,0.0013831458],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010748535,0.0010798118,0.2108339,0.0039730687,0.0009571384,0.006462177,0.005198299,0.106396586,0.0196657,0.023039617,0.043590866,0.577728],"study_design_scores_gemma":[0.00020742911,0.0005539224,0.058237497,0.0010959989,0.0010313401,0.0038838198,0.00357927,0.67145705,0.068974786,0.021439044,0.1692035,0.0003363845],"about_ca_topic_score_codex":0.008823761,"about_ca_topic_score_gemma":0.013760286,"teacher_disagreement_score":0.0138086835,"about_ca_system_score_codex":0.00168437,"about_ca_system_score_gemma":0.0035119187,"threshold_uncertainty_score":0.031975567},"labels":[],"label_agreement":null},{"id":"W3214740254","doi":"","title":"Software Migration: A Theoretical Framework (A Grounded Theory approach on Systematic Literature Review)","year":2021,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Berger (Canada)","funders":"","keywords":"Grounded theory; Systematic review; Computer science; Epistemology; Software; Sociology; Programming language; Qualitative research; Social science; Philosophy; Political science; MEDLINE","score_opus":0.011529987153419853,"score_gpt":0.24118195774340923,"score_spread":0.22965197058998937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214740254","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0195774,0.34829342,0.45217064,0.06703269,0.00255651,0.06581311,0.010346572,0.0006778995,0.033531815],"genre_scores_gemma":[0.106108174,0.11285623,0.68918365,0.00827137,0.0002805737,0.07809929,0.00338067,0.00008262051,0.0017374445],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"systematic_review","domain_scores_codex":[0.9177338,0.060303338,0.009292673,0.0043386607,0.006692595,0.0016389769],"domain_scores_gemma":[0.9165992,0.06879013,0.0041513196,0.0025157244,0.006914338,0.0010293454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08858309,0.0023340096,0.005147788,0.053093363,0.0046557067,0.014219591,0.00563986,0.0076836436,0.0053442726],"category_scores_gemma":[0.08222066,0.002284886,0.0040615182,0.041719466,0.010898323,0.0178719,0.01015339,0.0035706074,0.0011780991],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001263006,0.00015220395,0.0018826848,0.17874655,0.0007934392,0.0008066392,0.0317542,0.0038240673,0.0008886318,0.6055302,0.0149783725,0.16051663],"study_design_scores_gemma":[0.00043585495,0.0003187713,0.0022612538,0.33224085,0.0015043917,0.0007357203,0.055529043,0.0051044715,0.000747257,0.39466017,0.20624496,0.00021720014],"about_ca_topic_score_codex":0.011524433,"about_ca_topic_score_gemma":0.013758713,"teacher_disagreement_score":0.08858309,"about_ca_system_score_codex":0.024730595,"about_ca_system_score_gemma":0.07288713,"threshold_uncertainty_score":0.46847773},"labels":[],"label_agreement":null},{"id":"W3214949213","doi":"10.1109/tem.2021.3122012","title":"Toward Using Package Centrality Trend to Identify Packages in Decline","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Engineering Management","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Université Laval; Carleton University; Concordia University","funders":"","keywords":"Centrality; Computer science; Popularity; Reuse; Scalability; Code (set theory); Software; Software engineering; Code reuse; Data science; World Wide Web; Database; Operating system; Engineering","score_opus":0.03937547079648684,"score_gpt":0.30553573147891133,"score_spread":0.2661602606824245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214949213","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6799296,0.0009906288,0.27148858,0.0023028573,0.0002556135,0.0007601636,0.010642756,0.014607278,0.019022373],"genre_scores_gemma":[0.7025235,0.000315856,0.2770004,0.000296331,0.00016248303,0.0004939837,0.01381702,0.00086133025,0.0045291474],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952602,0.001044296,0.00040586607,0.0012879425,0.0016828082,0.00031892236],"domain_scores_gemma":[0.9666067,0.008570445,0.0062606507,0.003403969,0.01325926,0.0018991436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059524737,0.0016117878,0.0011094767,0.022860749,0.0011577532,0.004413697,0.0018907143,0.0016079257,0.0025388845],"category_scores_gemma":[0.03088207,0.0005162811,0.00096049876,0.012384845,0.0006789484,0.007242196,0.003947553,0.0016724238,0.0022893234],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028008042,0.0005332694,0.75528824,0.00031499012,0.00020341191,0.0002636368,0.002166335,0.008539515,0.0058237207,0.0058744843,0.015741805,0.20497058],"study_design_scores_gemma":[0.0001308072,0.0005342999,0.3887781,0.00019577044,0.00026638416,0.0007334383,0.005617654,0.52810436,0.0107683595,0.020196484,0.044389445,0.00028487295],"about_ca_topic_score_codex":0.018774988,"about_ca_topic_score_gemma":0.025641069,"teacher_disagreement_score":0.022860749,"about_ca_system_score_codex":0.0013807991,"about_ca_system_score_gemma":0.0025381579,"threshold_uncertainty_score":0.037331402},"labels":[],"label_agreement":null},{"id":"W3215087645","doi":"10.1155/2021/5069016","title":"A Novel Rank Aggregation‐Based Hybrid Multifilter Wrapper Feature Selection Method in Software Defect Prediction","year":2021,"lang":"en","type":"article","venue":"Computational Intelligence and Neuroscience","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Feature selection; Rank (graph theory); Computer science; Feature (linguistics); Software; Selection (genetic algorithm); Pattern recognition (psychology); Artificial intelligence; Data mining; Machine learning; Mathematics","score_opus":0.042036925436791335,"score_gpt":0.3140236752147824,"score_spread":0.2719867497779911,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3215087645","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036410697,0.00086898217,0.9606029,0.00015824058,0.000061929735,0.00006615317,0.00018266555,0.0010254132,0.0006230208],"genre_scores_gemma":[0.63136375,0.0007268916,0.3606735,0.00025140351,0.0001938493,0.00028578518,0.0013873446,0.00014260656,0.0049748244],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991387,0.0001380478,0.00006222065,0.00020672682,0.00035629017,0.000098042285],"domain_scores_gemma":[0.999321,0.00021297242,0.000075754004,0.00005919797,0.0002921411,0.00003898043],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010494731,0.0010478866,0.0013477756,0.0017985189,0.00047230284,0.00070440903,0.0010142868,0.00069778506,0.0010236118],"category_scores_gemma":[0.001998964,0.00030523745,0.0012664722,0.0014811575,0.0002752463,0.00093343604,0.0006277179,0.00063165784,0.00047639574],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002930592,0.00017705355,0.006275301,0.00015665086,0.00021512213,0.0002660574,0.00012684593,0.13895695,0.025360536,0.0021292095,0.006762594,0.8192806],"study_design_scores_gemma":[0.000023310176,0.00010514303,0.0026219387,0.000012791342,0.00005072645,0.000107491374,0.00002515062,0.9891661,0.005001477,0.0011296958,0.0017325627,0.00002367701],"about_ca_topic_score_codex":0.0045976066,"about_ca_topic_score_gemma":0.0039408165,"teacher_disagreement_score":0.0045976066,"about_ca_system_score_codex":0.00035373375,"about_ca_system_score_gemma":0.000872333,"threshold_uncertainty_score":0.009141684},"labels":[],"label_agreement":null},{"id":"W3215176293","doi":"10.1109/re51729.2021.00023","title":"Automated Traceability for Domain Modelling Decisions Empowered by Artificial Intelligence","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Traceability; Computer science; TRACE (psycholinguistics); Domain (mathematical analysis); Tracing; Set (abstract data type); Completeness (order theory); Artificial intelligence; Domain model; Software engineering; Data science; Machine learning; Domain knowledge; Programming language","score_opus":0.06684052606423473,"score_gpt":0.33346453587189595,"score_spread":0.2666240098076612,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3215176293","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010196112,0.00018073866,0.977471,0.0007237373,0.000031211464,0.00023257908,0.0002208671,0.00919034,0.0017535134],"genre_scores_gemma":[0.13970862,0.00033685094,0.8548354,0.00018569248,0.000025611229,0.00021879793,0.0017914859,0.0010274594,0.0018701196],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9843779,0.007328262,0.0012380992,0.0018524631,0.0047614933,0.00044176658],"domain_scores_gemma":[0.9540625,0.02497244,0.0029568358,0.014167753,0.0033464017,0.0004941312],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011666374,0.0020118083,0.0011994776,0.0060529616,0.0016531905,0.005975567,0.0036378028,0.0023834477,0.004581804],"category_scores_gemma":[0.05814569,0.0013508757,0.0028525856,0.0029762364,0.0021447162,0.010841414,0.0072018467,0.005334312,0.0015246158],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037120958,0.0009003106,0.005120311,0.0009766304,0.00020368054,0.0009060501,0.003997297,0.07253293,0.017515734,0.08480281,0.008389241,0.80428386],"study_design_scores_gemma":[0.000100328856,0.00018045437,0.001634285,0.0005115554,0.00012344167,0.00063855754,0.0008463887,0.73964185,0.03595948,0.16303745,0.057177436,0.00014880879],"about_ca_topic_score_codex":0.007123172,"about_ca_topic_score_gemma":0.008933804,"teacher_disagreement_score":0.011666374,"about_ca_system_score_codex":0.0024092933,"about_ca_system_score_gemma":0.00466981,"threshold_uncertainty_score":0.061698437},"labels":[],"label_agreement":null},{"id":"W3215908476","doi":"10.48550/arxiv.2112.01259","title":"Borrowing from Similar Code: A Deep Learning NLP-Based Approach for Log Statement Automation","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Source code; Artificial intelligence; Java; Context (archaeology); Natural language processing; Parsing; Machine learning; Information retrieval; Data mining; Programming language","score_opus":0.07292306413508969,"score_gpt":0.21502333995186096,"score_spread":0.14210027581677126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3215908476","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10679723,0.000679129,0.85175586,0.001406803,0.00013229366,0.00036919134,0.0034159422,0.031551737,0.0038918534],"genre_scores_gemma":[0.4560994,0.00029889378,0.5256983,0.0007885439,0.00008655296,0.0004040559,0.010428025,0.0007376595,0.0054584835],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989292,0.00018568661,0.00008845498,0.0004602665,0.0002418318,0.00009453035],"domain_scores_gemma":[0.99617,0.0019077647,0.00055126374,0.0006027155,0.00064091466,0.00012732144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009335841,0.0012937512,0.00048427723,0.0029806332,0.0005779179,0.0010242634,0.0019196568,0.0012227835,0.001562086],"category_scores_gemma":[0.0070722187,0.0004970898,0.0010086659,0.0018946724,0.0007838405,0.002987998,0.0016946654,0.0020993783,0.0012417572],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031745774,0.00065905706,0.023602419,0.00045352,0.000104190745,0.0010661453,0.0011987656,0.06416587,0.015817951,0.0047909073,0.019364385,0.8684593],"study_design_scores_gemma":[0.000032192653,0.00007200906,0.0029819629,0.000053023286,0.000044802317,0.00018550025,0.00018147155,0.96984416,0.007562671,0.013148483,0.005866007,0.00002773672],"about_ca_topic_score_codex":0.0122608645,"about_ca_topic_score_gemma":0.022810165,"teacher_disagreement_score":0.0122608645,"about_ca_system_score_codex":0.00089888053,"about_ca_system_score_gemma":0.0017334776,"threshold_uncertainty_score":0.024378955},"labels":[],"label_agreement":null},{"id":"W3216390711","doi":"10.32920/ryerson.14665455.v1","title":"Speeding up calibration of latent Dirichlet allocation model to improve topic analysis in software engineering","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Topic model; Exploit; Software; Simple (philosophy); Dirichlet distribution; Information retrieval; Data mining; Data science; Theoretical computer science; Mathematics; Programming language","score_opus":0.027622280103591152,"score_gpt":0.27322675240945776,"score_spread":0.2456044723058666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3216390711","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016074697,0.0004281268,0.9811916,0.0003521687,0.00007010012,0.00007861936,0.00010127336,0.0010220807,0.00068142585],"genre_scores_gemma":[0.2628559,0.0006363863,0.7319005,0.00034434898,0.00019912483,0.0006269001,0.0010826963,0.00066886825,0.0016852921],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9928925,0.004406427,0.00036641918,0.0012625335,0.00075831474,0.00031377963],"domain_scores_gemma":[0.97173667,0.0223188,0.00083304266,0.002463382,0.0023113696,0.00033677102],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011431054,0.001254036,0.0015961614,0.0031625964,0.0014487596,0.0031465702,0.0019021716,0.0021739078,0.0031682358],"category_scores_gemma":[0.06728162,0.0010946462,0.0015153886,0.0028709115,0.0012158736,0.0061109015,0.0040414217,0.0045800866,0.0023945302],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047744322,0.0003240684,0.010789215,0.0004557642,0.00025367012,0.00019449824,0.0026389915,0.41778693,0.011137599,0.058802728,0.01017349,0.4869657],"study_design_scores_gemma":[0.000036798996,0.00003084851,0.00095021335,0.000035988844,0.000022957398,0.000050861523,0.0001450951,0.9486023,0.002643851,0.044681683,0.002757736,0.000041715513],"about_ca_topic_score_codex":0.0068784696,"about_ca_topic_score_gemma":0.00745724,"teacher_disagreement_score":0.011431054,"about_ca_system_score_codex":0.0015563731,"about_ca_system_score_gemma":0.0025388482,"threshold_uncertainty_score":0.06045389},"labels":[],"label_agreement":null},{"id":"W3216427897","doi":"10.1109/scam52516.2021.00025","title":"PYREF: Refactoring Detection in Python Projects","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; Python (programming language); Computer science; Programming language; Java; Oracle; Software; Software engineering; Source code","score_opus":0.030539399267037633,"score_gpt":0.2773174201593533,"score_spread":0.24677802089231565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3216427897","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11770379,0.000774392,0.2154778,0.00050553726,0.00021147693,0.00048047127,0.025810266,0.6342174,0.0048189135],"genre_scores_gemma":[0.42220834,0.0005744655,0.45946023,0.00073117565,0.000074732234,0.0009856657,0.080754414,0.029783996,0.0054269815],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9928457,0.0011493559,0.001046413,0.0017396712,0.0026975486,0.0005213186],"domain_scores_gemma":[0.9861499,0.0062736003,0.0024375187,0.0027445515,0.0019785352,0.00041585282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006281947,0.001799956,0.0007405882,0.0034311395,0.00073091104,0.001706667,0.0025681648,0.0013207268,0.0032678796],"category_scores_gemma":[0.029264504,0.0010885653,0.0010837849,0.0020708248,0.0007988015,0.0037361085,0.0027844566,0.0017274488,0.003226141],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002154029,0.0007323139,0.16194248,0.0037235608,0.0004942827,0.0024955967,0.0025917464,0.023733785,0.05034222,0.007966527,0.2388568,0.50496674],"study_design_scores_gemma":[0.00047902024,0.0006686188,0.09043525,0.00065815804,0.00021807426,0.0037746516,0.0006170235,0.48945415,0.19510214,0.016090859,0.2019235,0.0005786375],"about_ca_topic_score_codex":0.0034617502,"about_ca_topic_score_gemma":0.003824482,"teacher_disagreement_score":0.006281947,"about_ca_system_score_codex":0.00071648776,"about_ca_system_score_gemma":0.0021413972,"threshold_uncertainty_score":0.033222497},"labels":[],"label_agreement":null},{"id":"W3216532022","doi":"10.32920/ryerson.14660433.v1","title":"Predicting the time-to-deliver of software changes","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software; Identification (biology); Data mining; Event (particle physics); Feature (linguistics); Data science","score_opus":0.02067907124096478,"score_gpt":0.25595874584334344,"score_spread":0.23527967460237867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3216532022","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91503507,0.00065464573,0.06849921,0.00087050174,0.00010864753,0.00014888844,0.011654763,0.0007768226,0.002251385],"genre_scores_gemma":[0.96289086,0.00042947047,0.022410752,0.000082634186,0.00007586889,0.00013629923,0.0117967455,0.00007808189,0.0020993652],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983041,0.00041763618,0.00016281834,0.0004594121,0.00050735974,0.0001486833],"domain_scores_gemma":[0.9397877,0.041068714,0.010319922,0.0037511927,0.0038037102,0.0012688002],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004773933,0.0006404128,0.0005134913,0.0031235341,0.0002640827,0.0012103948,0.00079939957,0.0013710303,0.0025711958],"category_scores_gemma":[0.055679746,0.00029651125,0.0006750394,0.0024200964,0.00030866897,0.001806582,0.0006568018,0.0016618052,0.0018155875],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004680788,0.0004776906,0.74086624,0.0003936827,0.000192316,0.00025571254,0.00056604174,0.108918995,0.0030330364,0.0053631915,0.004911493,0.13455355],"study_design_scores_gemma":[0.000044832454,0.0009049158,0.48454747,0.00014066728,0.00013157804,0.000389272,0.0006961826,0.485274,0.005954447,0.012241619,0.009570254,0.00010483655],"about_ca_topic_score_codex":0.0042277114,"about_ca_topic_score_gemma":0.0066036675,"teacher_disagreement_score":0.004773933,"about_ca_system_score_codex":0.0006780945,"about_ca_system_score_gemma":0.00063022634,"threshold_uncertainty_score":0.025247276},"labels":[],"label_agreement":null},{"id":"W3216616474","doi":"10.18280/isi.260504","title":"Reusable Component Retrieval from a Large Repository Using Word2Vec with Continuous Bag of Words","year":2021,"lang":"en","type":"article","venue":"Ingénierie des systèmes d information","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Word2vec; Component (thermodynamics); Computer science; Word embedding; Word (group theory); Code (set theory); Process (computing); Representation (politics); Embedding; Artificial intelligence; Information retrieval; Natural language processing; Data mining; Programming language","score_opus":0.012317461896000895,"score_gpt":0.2283517562152507,"score_spread":0.2160342943192498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3216616474","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20367225,0.004595357,0.74602896,0.0010159009,0.0013258782,0.0005147005,0.011988667,0.025982063,0.004876127],"genre_scores_gemma":[0.5537246,0.0026825657,0.35689273,0.00064223947,0.0005813232,0.00089049526,0.06716427,0.0015761459,0.015845587],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993092,0.00009308581,0.00007083648,0.0002069785,0.00021287914,0.000107018255],"domain_scores_gemma":[0.9994381,0.00014736937,0.00006125006,0.0001041245,0.00021945327,0.000029647954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00041247025,0.0023337677,0.001001252,0.0030188563,0.0005047,0.000882455,0.0009866257,0.0007603807,0.0028669813],"category_scores_gemma":[0.0016373113,0.00044538712,0.0013961922,0.0032756063,0.00038714215,0.0019215737,0.001180995,0.001044082,0.0029831324],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057273486,0.00038006302,0.0044659674,0.0008596936,0.0003802657,0.0010208734,0.0002911703,0.040600553,0.050713122,0.002558671,0.062924236,0.83523273],"study_design_scores_gemma":[0.00009463558,0.00049860793,0.0059577636,0.00009648066,0.00022940582,0.0011100384,0.00034931937,0.9181429,0.042979006,0.0069957483,0.023425013,0.00012108313],"about_ca_topic_score_codex":0.009806246,"about_ca_topic_score_gemma":0.016269598,"teacher_disagreement_score":0.009806246,"about_ca_system_score_codex":0.0005430412,"about_ca_system_score_gemma":0.001237645,"threshold_uncertainty_score":0.019498348},"labels":[],"label_agreement":null},{"id":"W3349379","doi":"10.1007/978-3-319-09940-8_9","title":"Bi-objective Genetic Search for Release Planning in Support of Themes","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Sorting; Theme (computing); Genetic algorithm; Set (abstract data type); Operations research; Multi-objective optimization; Preference; Management science; Machine learning; Algorithm; World Wide Web; Mathematics","score_opus":0.023381579251515135,"score_gpt":0.2838028446062741,"score_spread":0.26042126535475896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3349379","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034792107,0.0005324977,0.95338935,0.00019759222,0.000100765785,0.00011895404,0.00010169644,0.0005834744,0.010183512],"genre_scores_gemma":[0.503019,0.00037189425,0.48435345,0.00016089313,0.000065407556,0.00043246456,0.00033687247,0.00027599395,0.010984031],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999566,0.00014519732,0.000020441443,0.00006491812,0.00012631285,0.00007700804],"domain_scores_gemma":[0.9990237,0.0006539646,0.00006502394,0.000060453982,0.00014100007,0.000055836248],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013813976,0.0010854289,0.0012903136,0.0013296284,0.0006486956,0.0010234797,0.0018863506,0.0018953752,0.004982114],"category_scores_gemma":[0.0026065921,0.0007457759,0.0009538466,0.001406647,0.00066486915,0.00095548056,0.001176765,0.0014294827,0.0006647067],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006670501,0.000055123634,0.00018697287,0.000057658774,0.000033733086,0.000048805734,0.000044902128,0.9576373,0.0008740403,0.005535187,0.00092078943,0.034538712],"study_design_scores_gemma":[0.0000145335125,0.000030333023,0.00004492611,0.000007286162,0.000008489824,0.000009213718,0.000009441646,0.99783856,0.00020407839,0.0015798277,0.00024972798,0.000003590247],"about_ca_topic_score_codex":0.00605524,"about_ca_topic_score_gemma":0.0054419995,"teacher_disagreement_score":0.00605524,"about_ca_system_score_codex":0.00094821915,"about_ca_system_score_gemma":0.0014224803,"threshold_uncertainty_score":0.01666683},"labels":[],"label_agreement":null},{"id":"W33701851","doi":"","title":"A Software Design Pattern Based Approach to Auto Dynamic Difficulty in Video Games","year":2014,"lang":"en","type":"article","venue":"The Dental register","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Software; Computer graphics (images); Artificial intelligence; Software engineering; Programming language","score_opus":0.02231819023931951,"score_gpt":0.2548587565282449,"score_spread":0.2325405662889254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W33701851","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011884071,0.00013122777,0.98130643,0.0006477168,0.000038592818,0.00069841224,0.0001342843,0.0016993023,0.0034600056],"genre_scores_gemma":[0.03366697,0.00011208066,0.96361107,0.00009189106,0.000008494364,0.00047352086,0.00020390481,0.00018418295,0.0016479057],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99099594,0.0034057198,0.0012019782,0.0012606689,0.0028143625,0.0003213629],"domain_scores_gemma":[0.98918355,0.0050216964,0.0011924156,0.002084349,0.002209835,0.0003081998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066932156,0.0015071399,0.0005672038,0.003232984,0.0012750791,0.0038357743,0.002155261,0.0016131863,0.0019412182],"category_scores_gemma":[0.016988097,0.0013477742,0.0018673597,0.0018750349,0.0025951786,0.0032449116,0.0026276894,0.0029838984,0.0006937912],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027842144,0.0008792395,0.011957839,0.002216371,0.0002508831,0.0016591259,0.021923762,0.041315664,0.049918696,0.1802105,0.010370932,0.6790186],"study_design_scores_gemma":[0.0002852765,0.0013324567,0.007239476,0.0018887796,0.00048068914,0.00600079,0.0073438874,0.42649156,0.055194713,0.18361942,0.309757,0.00036599048],"about_ca_topic_score_codex":0.0051963395,"about_ca_topic_score_gemma":0.01010418,"teacher_disagreement_score":0.0066932156,"about_ca_system_score_codex":0.001842883,"about_ca_system_score_gemma":0.0038918436,"threshold_uncertainty_score":0.03539753},"labels":[],"label_agreement":null},{"id":"W349099688","doi":"","title":"An implementation for merging images for version control","year":2006,"lang":"en","type":"article","venue":"Annual Conference on Computers","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Maintainability; Scalability; Software; Merge (version control); Source code; Software engineering; Software development; Conflict resolution; Data mining; Programming language; Information retrieval; Database","score_opus":0.01805899312722866,"score_gpt":0.3203192417208621,"score_spread":0.3022602485936335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W349099688","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004964089,0.00019335195,0.852393,0.00023715026,0.000285574,0.0005824581,0.00047156657,0.13588384,0.004988996],"genre_scores_gemma":[0.067117095,0.00020159956,0.90146476,0.0003370047,0.00014882936,0.00092104985,0.001781261,0.0170921,0.010936295],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99542916,0.0006772646,0.00066766975,0.0010516132,0.0017265847,0.00044761444],"domain_scores_gemma":[0.9883837,0.0036122876,0.00086447055,0.003690334,0.0028769567,0.0005722761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005343907,0.0017494713,0.0010465049,0.0029970733,0.0013825584,0.0040974203,0.00424109,0.0020678972,0.025861071],"category_scores_gemma":[0.020214183,0.002073792,0.0015788332,0.002153542,0.001427644,0.0065193884,0.0037586852,0.0031399122,0.007959627],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014020522,0.00045611573,0.005403976,0.0010212602,0.00023019861,0.0008970212,0.0021603752,0.003417201,0.047255818,0.026827665,0.06546273,0.8454656],"study_design_scores_gemma":[0.0009399007,0.0010447974,0.005158088,0.0004979897,0.00048240862,0.0027843467,0.00057683315,0.1291315,0.3129905,0.030688358,0.5149861,0.00071919],"about_ca_topic_score_codex":0.0022365279,"about_ca_topic_score_gemma":0.0019857725,"teacher_disagreement_score":0.025861071,"about_ca_system_score_codex":0.0011623817,"about_ca_system_score_gemma":0.0017230166,"threshold_uncertainty_score":0.08651394},"labels":[],"label_agreement":null},{"id":"W37293407","doi":"10.71781/10886","title":"DECOR : détection et correction des défauts dans les systèmes orientés objet","year":2008,"lang":"fr","type":"dissertation","venue":"Frontiers in Psychiatry","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Code smell; Code refactoring; Computer science; Restructuring; Software engineering; Software maintenance; Software; Code (set theory); Source code; Software system; Abstraction; Software development; Java; Software design; Programming language; Software quality","score_opus":0.015033910986077311,"score_gpt":0.2811223211123236,"score_spread":0.2660884101262463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W37293407","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40529704,0.0027014816,0.5258078,0.0005831604,0.00066210795,0.00031627525,0.0024861854,0.0567744,0.0053715413],"genre_scores_gemma":[0.6554668,0.0006614823,0.32805264,0.00024981739,0.00009342608,0.00012952805,0.0025135772,0.0019523295,0.010880386],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981179,0.00022661044,0.00009552487,0.000462401,0.00096774986,0.00012984693],"domain_scores_gemma":[0.99667645,0.0017302162,0.00036196623,0.00047571486,0.00067543756,0.00008020863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001220586,0.0012688466,0.0006726528,0.001384198,0.00043089705,0.001187979,0.0007956693,0.001129861,0.003736737],"category_scores_gemma":[0.008055051,0.00047793722,0.00085492135,0.00051948754,0.00071563915,0.0011806446,0.0007549699,0.0007738674,0.0018008694],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011813589,0.00017181813,0.038458817,0.0007398102,0.000268953,0.0017732693,0.00072462903,0.021590738,0.116857104,0.001791732,0.011465646,0.8049761],"study_design_scores_gemma":[0.00021319037,0.0009925541,0.106555596,0.000189836,0.00026623573,0.0050760857,0.00072027004,0.54953647,0.3051129,0.0036384647,0.027497018,0.00020148493],"about_ca_topic_score_codex":0.007678726,"about_ca_topic_score_gemma":0.011162457,"teacher_disagreement_score":0.007678726,"about_ca_system_score_codex":0.000405973,"about_ca_system_score_gemma":0.0007418487,"threshold_uncertainty_score":0.015268028},"labels":[],"label_agreement":null},{"id":"W397180395","doi":"10.1007/978-1-84800-044-5_11","title":"Selecting Empirical Methods for Software Engineering Research","year":2007,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1135,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; National Research Council Canada; University of Toronto","funders":"","keywords":"Empirical research; Computer science; Variety (cybernetics); Data science; Software engineering; Software; Management science; Method engineering; Systems engineering; Engineering; Artificial intelligence; Mathematics","score_opus":0.24589088857868543,"score_gpt":0.48988522473861085,"score_spread":0.2439943361599254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W397180395","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024271298,0.034404863,0.87819064,0.0060424358,0.0014667556,0.00049922935,0.0004151573,0.0008944089,0.07565932],"genre_scores_gemma":[0.026601594,0.032750715,0.89663255,0.0017199833,0.0012152232,0.0016917599,0.00071994954,0.001111841,0.03755639],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98060715,0.011877193,0.001031754,0.00068233476,0.005569965,0.00023161776],"domain_scores_gemma":[0.91617244,0.07307695,0.0012712164,0.003609219,0.0053413515,0.00052883365],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.02266292,0.0012262985,0.0016431717,0.008817729,0.0012454364,0.0062953373,0.002028918,0.0019254669,0.013276712],"category_scores_gemma":[0.06790333,0.0011725331,0.0009952973,0.009113512,0.0031947328,0.008017119,0.00217021,0.0038186999,0.0057358374],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037967933,0.00008030939,0.0010269744,0.0016942744,0.00006342595,0.000086612585,0.00080659124,0.0012119107,0.00059656374,0.40379214,0.05601337,0.5345899],"study_design_scores_gemma":[0.00005302151,0.000063666965,0.0015921418,0.0035862878,0.00008311977,0.0003731442,0.00096229,0.007803162,0.0014994843,0.72647727,0.2574428,0.00006364456],"about_ca_topic_score_codex":0.00072258705,"about_ca_topic_score_gemma":0.0018311501,"teacher_disagreement_score":0.97733706,"about_ca_system_score_codex":0.002098207,"about_ca_system_score_gemma":0.002402486,"threshold_uncertainty_score":0.11985445},"labels":[],"label_agreement":null},{"id":"W40894950","doi":"","title":"FiGD: An Open Source Intellectual Property Violation Detector.","year":2009,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Code refactoring; Computer science; Java; Open source; Obfuscation; Fingerprint (computing); Detector; Property (philosophy); Generator (circuit theory); Intellectual property; Source code; Open source software; Programming language; Software; Operating system; Computer security","score_opus":0.019583414232368646,"score_gpt":0.2584719341747957,"score_spread":0.23888851994242705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W40894950","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0260209,0.000806246,0.30632904,0.00047595263,0.00039060134,0.00036432987,0.018080698,0.64023995,0.007292344],"genre_scores_gemma":[0.35666567,0.0005730964,0.50467515,0.001208595,0.00017560979,0.00071294024,0.061286386,0.040653646,0.03404894],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984664,0.00017789961,0.000098346085,0.00028141457,0.00088175485,0.000094276824],"domain_scores_gemma":[0.99573,0.0017097559,0.00054934307,0.0010577186,0.0007431832,0.0002100577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012855454,0.0012975347,0.0007323477,0.0032048847,0.0004081581,0.0011401854,0.0023197322,0.0018674689,0.01276529],"category_scores_gemma":[0.0074069044,0.0008689699,0.0006934154,0.0011278823,0.00054830575,0.0022157878,0.0018440745,0.0010882304,0.008266644],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016466022,0.00036350038,0.018865163,0.0008906062,0.00021016173,0.0012204703,0.0002075206,0.0064376434,0.04902557,0.005452384,0.3443577,0.57132274],"study_design_scores_gemma":[0.0007300202,0.0007581047,0.016852522,0.00025492202,0.00014873045,0.0045844764,0.00014059174,0.41160706,0.3097601,0.015270942,0.23946194,0.0004306023],"about_ca_topic_score_codex":0.0011545445,"about_ca_topic_score_gemma":0.0013527592,"teacher_disagreement_score":0.01276529,"about_ca_system_score_codex":0.0006450419,"about_ca_system_score_gemma":0.00084271905,"threshold_uncertainty_score":0.042704165},"labels":[],"label_agreement":null},{"id":"W41831425","doi":"","title":"Variability Modeling: State of the Art and Future Directions.","year":2010,"lang":"en","type":"article","venue":"Variability Modelling of Software-Intensive Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Data science","score_opus":0.02096445860599052,"score_gpt":0.2344570107138745,"score_spread":0.21349255210788398,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W41831425","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0047269203,0.058264263,0.9303342,0.002231585,0.0002972629,0.00003422537,0.0001825004,0.00079538283,0.0031336718],"genre_scores_gemma":[0.42526788,0.18782993,0.37598446,0.0010752712,0.003127363,0.0003050222,0.0015816392,0.0006492418,0.004179233],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9976399,0.0009667022,0.00015470147,0.00045596238,0.00068248145,0.000100150326],"domain_scores_gemma":[0.99053836,0.006631748,0.00050633564,0.0011859911,0.00095144415,0.00018608512],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056039663,0.0013147577,0.002179237,0.0015145988,0.00047071246,0.0047510834,0.0039622663,0.0018927167,0.0029149903],"category_scores_gemma":[0.01104558,0.0006465554,0.0017615117,0.0027153175,0.0013226665,0.005801017,0.001740987,0.0023388024,0.0011752923],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017745834,0.00024136428,0.006941679,0.0019624936,0.0005254403,0.00016219458,0.00029294676,0.2516923,0.0044010775,0.12675008,0.010273611,0.5965794],"study_design_scores_gemma":[0.000021810576,0.0000961791,0.0014000384,0.000567915,0.0001744326,0.00021013839,0.00018810833,0.67402077,0.002078508,0.284415,0.03672771,0.00009945242],"about_ca_topic_score_codex":0.0031796093,"about_ca_topic_score_gemma":0.0020619554,"teacher_disagreement_score":0.0056039663,"about_ca_system_score_codex":0.0008807976,"about_ca_system_score_gemma":0.0015058676,"threshold_uncertainty_score":0.02963698},"labels":[],"label_agreement":null},{"id":"W4200148139","doi":"10.1109/models-c53483.2021.00126","title":"Metamodel Refactoring using Constraint Solving: a Quality-based Perspective","year":2021,"lang":"en","type":"article","venue":"2021 ACM/IEEE International Conference on Model Driven Engineering Languages and Systems Companion (MODELS-C)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Code refactoring; Metamodeling; Computer science; Correctness; Software engineering; Quality (philosophy); Set (abstract data type); Task (project management); Programming language; Software; Systems engineering; Engineering","score_opus":0.1685616600302603,"score_gpt":0.3734979495000038,"score_spread":0.20493628946974352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200148139","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017029494,0.00028044215,0.99601483,0.0002992263,0.000015300884,0.00007053458,0.00004788209,0.00014630194,0.0014226083],"genre_scores_gemma":[0.026037565,0.00063723116,0.9720222,0.000092981965,0.000024618179,0.0001193772,0.00019485627,0.00013878617,0.00073235366],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945867,0.0024546522,0.0003570369,0.00059842144,0.0017251918,0.00027803675],"domain_scores_gemma":[0.9906222,0.006587946,0.0006557864,0.00092976366,0.0010190287,0.0001852456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052613914,0.002008703,0.0012086573,0.0023482498,0.00094290747,0.0032260227,0.003726179,0.0019220428,0.0028903438],"category_scores_gemma":[0.015242008,0.0010842445,0.0026396636,0.0033284493,0.0022789151,0.0027138146,0.0027015102,0.0030921425,0.000499189],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010480304,0.0002665611,0.0013409199,0.0013232525,0.00027715237,0.00057077536,0.00064996345,0.54299706,0.013651418,0.21651322,0.0036344633,0.21867046],"study_design_scores_gemma":[0.00008308077,0.00009498054,0.00025917534,0.00026811106,0.00011612944,0.00032412895,0.00020453653,0.853693,0.014690208,0.10858367,0.021619042,0.000063860636],"about_ca_topic_score_codex":0.007818611,"about_ca_topic_score_gemma":0.007713182,"teacher_disagreement_score":0.007818611,"about_ca_system_score_codex":0.0015794203,"about_ca_system_score_gemma":0.0036952589,"threshold_uncertainty_score":0.027825236},"labels":[],"label_agreement":null},{"id":"W4200448782","doi":"10.1007/s10664-021-10076-4","title":"Using code reviews to automatically configure static analysis tools","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Leverage (statistics); Code review; Static program analysis; Java; Code (set theory); Source code; Precision and recall; Context (archaeology); Information retrieval; Software engineering; Statement (logic); Programming language; Artificial intelligence; Software; Software development","score_opus":0.09646917508994118,"score_gpt":0.3629266213116876,"score_spread":0.26645744622174644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200448782","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43190795,0.001980081,0.38422498,0.0012347674,0.0005324467,0.0012065603,0.0035547805,0.1573122,0.01804623],"genre_scores_gemma":[0.69224,0.00037115795,0.28646386,0.0002745926,0.00016100264,0.0005091701,0.005263946,0.0074126185,0.0073036146],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9886114,0.0037619763,0.00089229766,0.0021777104,0.004205462,0.00035123108],"domain_scores_gemma":[0.8717375,0.06693242,0.016523622,0.013710327,0.028865745,0.002230371],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063219476,0.001538462,0.0010644357,0.009929089,0.0009738647,0.002639444,0.001645864,0.0010762541,0.0038429638],"category_scores_gemma":[0.08582751,0.0010322994,0.0006241817,0.0028378936,0.0003782726,0.0025604372,0.0017988621,0.001008924,0.0036700617],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056595076,0.00041241408,0.041433558,0.0007357438,0.00020370302,0.0005077586,0.0017176459,0.007318327,0.032192867,0.0018343468,0.038972877,0.8741047],"study_design_scores_gemma":[0.0006662958,0.0014759668,0.09610696,0.00098059,0.000689658,0.0019343899,0.002100208,0.67960757,0.1081542,0.011643512,0.096043296,0.0005973235],"about_ca_topic_score_codex":0.004487311,"about_ca_topic_score_gemma":0.0111260675,"teacher_disagreement_score":0.009929089,"about_ca_system_score_codex":0.0011304801,"about_ca_system_score_gemma":0.0031409257,"threshold_uncertainty_score":0.033434093},"labels":[],"label_agreement":null},{"id":"W4200576845","doi":"10.1109/models-c53483.2021.00090","title":"DoMoBOT: An AI-Empowered Bot for Automated and Interactive Domain Modelling","year":2021,"lang":"en","type":"article","venue":"2021 ACM/IEEE International Conference on Model Driven Engineering Languages and Systems Companion (MODELS-C)","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Domain (mathematical analysis); Set (abstract data type); Domain model; MODELLER; Artificial intelligence; Software engineering; Domain engineering; Machine learning; Programming language; Domain knowledge; Software; Software development","score_opus":0.07182052063672636,"score_gpt":0.340547473686961,"score_spread":0.26872695305023464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200576845","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004384261,0.00009631205,0.9545845,0.00036912822,0.000065512766,0.00029873592,0.00050653465,0.034739833,0.0049551483],"genre_scores_gemma":[0.07114591,0.00030532994,0.9113684,0.00048960693,0.000021380563,0.0004997078,0.0024817223,0.0047107297,0.008977091],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998069,0.0006533543,0.00016210161,0.00030096126,0.0006743351,0.0001403262],"domain_scores_gemma":[0.9962321,0.001929808,0.00022081617,0.0009932243,0.00037987658,0.00024422107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031163504,0.0009827736,0.0006095954,0.0018398383,0.00080748816,0.0024979513,0.0028230047,0.001948918,0.007268828],"category_scores_gemma":[0.008915739,0.0010498291,0.0015414709,0.00069688587,0.001444068,0.004377172,0.0055413917,0.0028589065,0.0029348054],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001281342,0.0011257071,0.0076903696,0.0020641005,0.0004430235,0.0020439683,0.0042915996,0.11728525,0.058422923,0.26237398,0.09852261,0.44445503],"study_design_scores_gemma":[0.000112169,0.00012254309,0.0004740369,0.00022418112,0.000052240746,0.0005261536,0.0002481072,0.7016674,0.015167195,0.065278925,0.21602927,0.00009774461],"about_ca_topic_score_codex":0.004196346,"about_ca_topic_score_gemma":0.009171707,"teacher_disagreement_score":0.007268828,"about_ca_system_score_codex":0.001398868,"about_ca_system_score_gemma":0.0025267906,"threshold_uncertainty_score":0.024316609},"labels":[],"label_agreement":null},{"id":"W4205290156","doi":"10.22215/etd/2021-14656","title":"Understanding How Developers Reuse Stack Overflow Code in Their GitHub Projects","year":2021,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Code reuse; Computer science; Source code; Reuse; Code (set theory); Software; Code review; Software evolution; Software engineering; Static program analysis; Programming language; Software development; Software construction; Engineering","score_opus":0.12211958829014888,"score_gpt":0.29806519202736353,"score_spread":0.17594560373721466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205290156","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95573866,0.00073390245,0.019731632,0.0029426608,0.000023447394,0.00008090828,0.00024166478,0.0006213871,0.019885687],"genre_scores_gemma":[0.972494,0.0010797875,0.01828415,0.00047381307,0.00001102955,0.00005300039,0.0004850205,0.00043650807,0.0066827303],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9967127,0.001216156,0.00014448077,0.0005119393,0.001028974,0.00038576446],"domain_scores_gemma":[0.9816295,0.009948087,0.0035807034,0.0014049093,0.0027910238,0.00064581836],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040503447,0.00036390842,0.00021956029,0.0030450919,0.001411041,0.0046948413,0.00069667905,0.001098596,0.0018623077],"category_scores_gemma":[0.03293705,0.0004922276,0.00031553104,0.00301205,0.0017992596,0.012917974,0.0026820556,0.0017344812,0.0006731325],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012197581,0.00032120198,0.35985738,0.0005238772,0.00008003588,0.0008589661,0.34057346,0.0015544026,0.010373552,0.014672666,0.009574207,0.2614883],"study_design_scores_gemma":[0.00003771921,0.00031186178,0.6163706,0.0012500838,0.00020471498,0.0019541632,0.23204248,0.012600757,0.0107168695,0.02329542,0.100980766,0.0002345413],"about_ca_topic_score_codex":0.017772853,"about_ca_topic_score_gemma":0.037932944,"teacher_disagreement_score":0.017772853,"about_ca_system_score_codex":0.0019403535,"about_ca_system_score_gemma":0.0027213672,"threshold_uncertainty_score":0.03533882},"labels":[],"label_agreement":null},{"id":"W4205420848","doi":"10.1007/s10515-021-00319-5","title":"Improving the prediction of continuous integration build failures using deep learning","year":2022,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Benchmark (surveying); Feature engineering; Process (computing); Machine learning; Artificial intelligence; Task (project management); Construct (python library); Outcome (game theory); Recurrent neural network; Deep learning; Software; Code (set theory); Feature (linguistics); Artificial neural network; Data mining; Engineering; Systems engineering","score_opus":0.008641220931652826,"score_gpt":0.22069118800838491,"score_spread":0.2120499670767321,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205420848","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73325413,0.0017818193,0.2563053,0.000755304,0.00020305812,0.000037422527,0.0008616496,0.0044111824,0.002390036],"genre_scores_gemma":[0.9871076,0.00011250769,0.010829752,0.00005779697,0.000031473064,0.000011564864,0.0006324796,0.000043718654,0.001173216],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995561,0.00007301444,0.0000243555,0.00013311578,0.00011928621,0.00009415318],"domain_scores_gemma":[0.9972017,0.0014523115,0.00032066458,0.00023199675,0.0006232563,0.00017013209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00085914327,0.00128294,0.0007417723,0.0011507,0.0002449366,0.00071605353,0.0010981833,0.0010312619,0.0012356938],"category_scores_gemma":[0.0038369591,0.0004696681,0.0004930637,0.0006845973,0.0003226416,0.0012712514,0.00078432355,0.001915004,0.00060401985],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045974695,0.00061388896,0.03538254,0.00009252821,0.00011900565,0.00020141595,0.000060077673,0.7616726,0.004627315,0.00095299724,0.006144384,0.1896735],"study_design_scores_gemma":[0.0000023201942,0.000013244291,0.0005595333,0.0000023499272,0.0000032196365,0.0000045474835,0.0000025006732,0.99867165,0.0003450783,0.00034082666,0.00005318157,0.0000015429404],"about_ca_topic_score_codex":0.011219354,"about_ca_topic_score_gemma":0.014653042,"teacher_disagreement_score":0.011219354,"about_ca_system_score_codex":0.000669302,"about_ca_system_score_gemma":0.0008945202,"threshold_uncertainty_score":0.022308111},"labels":[],"label_agreement":null},{"id":"W4205513494","doi":"10.1109/ase51524.2021.9678640","title":"Is Historical Data an Appropriate Benchmark for Reviewer Recommendation Systems? : A Case Study of the Gerrit Community","year":2021,"lang":"en","type":"article","venue":"2021 36th IEEE/ACM International Conference on Automated Software Engineering (ASE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Polytechnique Montréal; McGill University","funders":"","keywords":"Context (archaeology); Pessimism; Computer science; Benchmark (surveying); Replicate; Task (project management); Data science; History; Epistemology; Management","score_opus":0.1459378586426111,"score_gpt":0.3631536798642139,"score_spread":0.21721582122160277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205513494","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9839527,0.0011480495,0.0072677666,0.0018040782,0.00005674351,0.00043698473,0.0005413618,0.00026592854,0.0045263325],"genre_scores_gemma":[0.98261225,0.0003752549,0.014023932,0.00031556145,0.00007086336,0.0002358108,0.0006293292,0.00012573601,0.0016111885],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.95386857,0.030837,0.0025901757,0.0037093388,0.008021301,0.00097354257],"domain_scores_gemma":[0.60254973,0.27839342,0.02805292,0.020271076,0.06323623,0.0074966024],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.039902847,0.0004654721,0.0007826616,0.004814048,0.0038229723,0.0033390282,0.0019040874,0.0020959128,0.0012943348],"category_scores_gemma":[0.19410795,0.00047927347,0.00055477093,0.0048914477,0.0014764583,0.0049091033,0.0018036686,0.0014910189,0.0007009175],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001634384,0.0019466062,0.45260215,0.0028609708,0.00036229112,0.0071781348,0.18413486,0.00649564,0.009586056,0.0031119406,0.015753541,0.31433347],"study_design_scores_gemma":[0.00051459274,0.005599404,0.56155044,0.002055454,0.00045276823,0.0071232314,0.20917334,0.06415552,0.0165475,0.005892315,0.12601088,0.0009246325],"about_ca_topic_score_codex":0.020884013,"about_ca_topic_score_gemma":0.04176582,"teacher_disagreement_score":0.96009713,"about_ca_system_score_codex":0.0029612037,"about_ca_system_score_gemma":0.002903485,"threshold_uncertainty_score":0.211029},"labels":[],"label_agreement":null},{"id":"W4205795599","doi":"10.1109/ase51524.2021.9678554","title":"Automatically Annotating Sentences for Task-specific Bug Report Summarization","year":2021,"lang":"en","type":"article","venue":"2021 36th IEEE/ACM International Conference on Automated Software Engineering (ASE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge; University of Alberta","funders":"","keywords":"Automatic summarization; Computer science; Task (project management); DevOps; Annotation; Software engineering; Information retrieval; World Wide Web; Natural language processing; Software; Artificial intelligence; Programming language; Engineering","score_opus":0.03620246813544579,"score_gpt":0.3043772340185386,"score_spread":0.2681747658830928,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205795599","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1024886,0.0030717568,0.7949791,0.0038748542,0.0021543233,0.0014136005,0.03194063,0.048880737,0.011196372],"genre_scores_gemma":[0.14510994,0.0010763622,0.7696705,0.00071287935,0.0009311578,0.0014815573,0.06965651,0.0045882915,0.00677276],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9909019,0.00417686,0.0012998725,0.0015033435,0.0017301923,0.00038786232],"domain_scores_gemma":[0.9471923,0.021772387,0.006037434,0.004130394,0.019840287,0.0010272377],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009600126,0.0023829082,0.0012147884,0.008525165,0.0014664857,0.0027368965,0.0015674857,0.0016336404,0.004272044],"category_scores_gemma":[0.045108963,0.0008044557,0.0010476869,0.0040256507,0.0005918401,0.0037960478,0.0026523543,0.0020380586,0.005055931],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005924125,0.00031661207,0.015494099,0.004914821,0.0002792693,0.0012761202,0.011063268,0.0063450597,0.109378435,0.0065516075,0.27659723,0.56719106],"study_design_scores_gemma":[0.0002338581,0.0011421167,0.048833117,0.0016211713,0.0011536417,0.0021090005,0.00940651,0.18965702,0.14073984,0.028025704,0.57642996,0.0006481782],"about_ca_topic_score_codex":0.0024613836,"about_ca_topic_score_gemma":0.004875251,"teacher_disagreement_score":0.009600126,"about_ca_system_score_codex":0.0008242657,"about_ca_system_score_gemma":0.0024738477,"threshold_uncertainty_score":0.05077094},"labels":[],"label_agreement":null},{"id":"W4206030968","doi":"10.1007/s10270-021-00942-6","title":"Automated, interactive, and traceable domain modelling empowered by artificial intelligence","year":2022,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Rotation formalisms in three dimensions; Domain (mathematical analysis); Software engineering; Domain engineering; Domain analysis; Domain model; Software; Artificial intelligence; Domain knowledge; Software development; Programming language; Component-based software engineering; Software construction","score_opus":0.03171702711815094,"score_gpt":0.2704622975187825,"score_spread":0.23874527040063154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206030968","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062777675,0.00007801528,0.9858766,0.0002246164,0.000022014205,0.00005267739,0.00016998261,0.0037005714,0.003597746],"genre_scores_gemma":[0.25360948,0.00050175004,0.738555,0.00015454288,0.000026561693,0.00013506466,0.0012245349,0.00090605655,0.0048870696],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978077,0.000799059,0.00015846474,0.0003663795,0.00073417026,0.00013428606],"domain_scores_gemma":[0.9961862,0.0015681183,0.00021406579,0.0015723128,0.00034318192,0.00011607087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018732037,0.00078643596,0.0006711163,0.0014711902,0.000683863,0.0038295174,0.0021644125,0.0010793245,0.0046255495],"category_scores_gemma":[0.0077206087,0.0007017202,0.0014047901,0.00106869,0.0015854917,0.004901462,0.0046496447,0.002801348,0.0015982311],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032156782,0.00042188747,0.0036456694,0.0005223226,0.00013711952,0.0007450991,0.002202225,0.36661568,0.0229999,0.282064,0.007802679,0.3125219],"study_design_scores_gemma":[0.000022248076,0.000036244295,0.00026283896,0.00007245773,0.000038802515,0.0001656636,0.00017528477,0.8193854,0.011871503,0.14285377,0.02508196,0.000033841636],"about_ca_topic_score_codex":0.00635027,"about_ca_topic_score_gemma":0.009934974,"teacher_disagreement_score":0.00635027,"about_ca_system_score_codex":0.00090197945,"about_ca_system_score_gemma":0.002657458,"threshold_uncertainty_score":0.015474021},"labels":[],"label_agreement":null},{"id":"W4206541176","doi":"10.1109/tdsc.2021.3138700","title":"Dataset Characteristics for Reliable Code Authorship Attribution","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Dependable and Secure Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Attribution; Computer science; Coding (social sciences); Source code; Field (mathematics); Robustness (evolution); Code (set theory); Data science; Data mining; Benchmark (surveying); Information retrieval; Set (abstract data type); Artificial intelligence; Programming language","score_opus":0.03607006238654237,"score_gpt":0.2916618682774196,"score_spread":0.2555918058908772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206541176","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6087066,0.0028931967,0.15509069,0.00595029,0.0015013918,0.0040782522,0.19981243,0.0072114444,0.014755698],"genre_scores_gemma":[0.6287972,0.00044594705,0.14701244,0.0005116315,0.0003250233,0.004613325,0.21647654,0.0005835692,0.001234369],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.97141695,0.00794027,0.006367421,0.004792903,0.008388028,0.0010944877],"domain_scores_gemma":[0.85936826,0.0663571,0.010935705,0.03488844,0.026618404,0.0018321038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021742862,0.0010236684,0.0008972662,0.005259944,0.001911995,0.0044279303,0.0020824752,0.002235132,0.0020713215],"category_scores_gemma":[0.16109192,0.00025794443,0.0011331517,0.0066255946,0.0011924919,0.003644193,0.0024533921,0.0019246946,0.0015435456],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022594437,0.0016397886,0.4370867,0.0033818048,0.00067139056,0.00092147687,0.002251027,0.024679588,0.022388762,0.017431205,0.14222927,0.34505957],"study_design_scores_gemma":[0.00082164176,0.0012583083,0.4223519,0.0016447832,0.0006689233,0.0032036654,0.004394999,0.17645021,0.06071302,0.053467646,0.274552,0.0004729857],"about_ca_topic_score_codex":0.0018509376,"about_ca_topic_score_gemma":0.0022613497,"teacher_disagreement_score":0.021742862,"about_ca_system_score_codex":0.0015140689,"about_ca_system_score_gemma":0.002771867,"threshold_uncertainty_score":0.114988625},"labels":[],"label_agreement":null},{"id":"W4206639121","doi":"10.22215/etd/2021-14643","title":"Using Machine Learning to Detect Architectural Integrity Violations Associated with Bugs","year":2021,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Maintainability; Computer science; Architectural pattern; Software engineering; Architecture; Set (abstract data type); Software bug; Software; Debugging; Programming language; Software development; Software design","score_opus":0.029663435513647297,"score_gpt":0.3073713762723838,"score_spread":0.2777079407587365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206639121","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6762309,0.0024589377,0.30877888,0.0017492251,0.00032396917,0.0002091792,0.0008925025,0.003609265,0.0057470542],"genre_scores_gemma":[0.9285975,0.0004277883,0.06668098,0.00020049879,0.00013619369,0.000086609776,0.0015581816,0.00006518139,0.0022471834],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987191,0.0003416568,0.00011114643,0.00032060352,0.00033321025,0.00017417417],"domain_scores_gemma":[0.99230653,0.0049875034,0.00092996436,0.00050455314,0.0011033006,0.0001681636],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001996309,0.000814329,0.0008218489,0.0024096796,0.00053043873,0.0011266979,0.0009548817,0.0011170561,0.00097563764],"category_scores_gemma":[0.009416867,0.00027297268,0.0007377625,0.0013746931,0.000447633,0.0010600191,0.0006259629,0.0016048775,0.0005778345],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004535175,0.0011326293,0.11385149,0.00019375895,0.00032370444,0.0003499479,0.0002222733,0.14822812,0.008435897,0.0013945717,0.007265185,0.71814895],"study_design_scores_gemma":[0.000021917864,0.00016782724,0.011728024,0.000028072647,0.00006109683,0.00009223252,0.00006553745,0.97969925,0.0035565929,0.0037938852,0.00076450943,0.000021023638],"about_ca_topic_score_codex":0.0042553185,"about_ca_topic_score_gemma":0.004157243,"teacher_disagreement_score":0.0042553185,"about_ca_system_score_codex":0.00065132004,"about_ca_system_score_gemma":0.00069163047,"threshold_uncertainty_score":0.010557592},"labels":[],"label_agreement":null},{"id":"W4206923519","doi":"10.1142/s0218194021400192","title":"Automatically Generating Release Notes with Content Classification Models","year":2021,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Automatic summarization; Computer science; Task (project management); Metric (unit); Software; Software release life cycle; Artificial intelligence; Word (group theory); Machine learning; Word embedding; Natural language processing; Information retrieval; Embedding; Software development; Software quality; Programming language","score_opus":0.03481585465653141,"score_gpt":0.25578651031495636,"score_spread":0.22097065565842494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206923519","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17725454,0.0016201508,0.7849579,0.0006764891,0.00039814546,0.00061480474,0.005197963,0.026504321,0.002775726],"genre_scores_gemma":[0.42407382,0.0006548456,0.5480034,0.00020215659,0.0003755745,0.00047547123,0.019843949,0.00088593847,0.005484851],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988084,0.00023980348,0.00015151457,0.00038844143,0.00032432293,0.00008744272],"domain_scores_gemma":[0.9910596,0.0050878646,0.0010633433,0.0008445414,0.0017566294,0.00018800332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001524713,0.001492214,0.0007775201,0.0039564576,0.0004306827,0.0011488168,0.0014301948,0.0009256172,0.0012314662],"category_scores_gemma":[0.007787051,0.0004578079,0.0010188358,0.0019141454,0.00036388362,0.0022635746,0.00072130805,0.0012129131,0.0015356796],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000746841,0.0006218923,0.010703673,0.00060952717,0.00012558981,0.0007407578,0.00043468812,0.06555104,0.029526953,0.0027791653,0.026589012,0.86157084],"study_design_scores_gemma":[0.00007805935,0.00018233067,0.0031530906,0.000037595997,0.000102621176,0.00015683955,0.00011005886,0.9669115,0.020346384,0.0037442413,0.005142037,0.000035276367],"about_ca_topic_score_codex":0.003488599,"about_ca_topic_score_gemma":0.004628033,"teacher_disagreement_score":0.0039564576,"about_ca_system_score_codex":0.0008055131,"about_ca_system_score_gemma":0.0010735175,"threshold_uncertainty_score":0.008063555},"labels":[],"label_agreement":null},{"id":"W4206994298","doi":"10.1109/asew52652.2021.00017","title":"Toward a Smell-aware Prediction Model for CI Build Failures","year":2021,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code smell; Computer science; Context (archaeology); Code (set theory); Quality (philosophy); Predictive modelling; Detector; Artificial intelligence; Reliability engineering; Machine learning; Software quality; Data mining; Engineering; Software; Programming language; Telecommunications; Software development","score_opus":0.03745704801880494,"score_gpt":0.2782276312020863,"score_spread":0.24077058318328137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206994298","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5557884,0.0030940874,0.40307572,0.0016208534,0.00022618719,0.00024778818,0.011073354,0.022214096,0.0026595497],"genre_scores_gemma":[0.8921537,0.00043084144,0.092533395,0.00022492741,0.00012355177,0.00016931212,0.012543631,0.00033340362,0.0014872147],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99840134,0.00030619503,0.000153437,0.0005752538,0.00039185505,0.00017198375],"domain_scores_gemma":[0.99118465,0.004480536,0.001369869,0.0006895982,0.0018584612,0.00041682814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027285484,0.002202827,0.0012262838,0.0057213292,0.00036554923,0.0015560531,0.0014841612,0.0013929927,0.0007479692],"category_scores_gemma":[0.010825759,0.00053631165,0.0010552447,0.0021289526,0.00036468898,0.0018549593,0.0011791431,0.001925509,0.001514601],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006587377,0.000958322,0.39019996,0.00064490124,0.0004057787,0.0006687066,0.0004109152,0.26642302,0.012948343,0.0011451885,0.023403669,0.30213252],"study_design_scores_gemma":[0.000011681728,0.00010279181,0.017209614,0.000036410853,0.000046226196,0.00013273343,0.000051204446,0.978049,0.0020497518,0.001171565,0.0011172444,0.00002180277],"about_ca_topic_score_codex":0.009585964,"about_ca_topic_score_gemma":0.01053746,"teacher_disagreement_score":0.009585964,"about_ca_system_score_codex":0.00071945146,"about_ca_system_score_gemma":0.001215902,"threshold_uncertainty_score":0.019060373},"labels":[],"label_agreement":null},{"id":"W4210265361","doi":"10.32920/19067156.v1","title":"Simulation Methods and Software Engineering Principles for the Analysis of Cost-Effective Disease Prevention","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Reuse; Software engineering; Health care; Reusability; Software; Medicine; Computer science; Engineering; Operating system; Political science","score_opus":0.06371809419479925,"score_gpt":0.41071753546589634,"score_spread":0.3469994412710971,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210265361","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024072572,0.0007742349,0.991859,0.00086984807,0.00008981162,0.00014138909,0.000078686244,0.00015056173,0.0036292628],"genre_scores_gemma":[0.11516434,0.003285705,0.8753918,0.000405785,0.00019710348,0.001682179,0.00026258753,0.0002568957,0.0033536716],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99402,0.0041706744,0.00028167668,0.0002971968,0.001071837,0.00015869703],"domain_scores_gemma":[0.96379507,0.0324579,0.0009588502,0.0012726577,0.0012592745,0.00025631508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01162926,0.0019998243,0.0016668848,0.0038051142,0.0008389597,0.0025378896,0.0022526574,0.0019220756,0.0054980726],"category_scores_gemma":[0.032437313,0.0012819525,0.0029243617,0.002349708,0.0036464669,0.0025637567,0.0021391576,0.0040331674,0.00062746095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028027789,0.000109250715,0.0013049892,0.00031084722,0.00023973337,0.00006020264,0.00016677327,0.48191822,0.000320336,0.48857203,0.0016761636,0.025293354],"study_design_scores_gemma":[0.00005512699,0.00007530627,0.00031109998,0.00018649007,0.00005773062,0.000029442026,0.00006946298,0.6376757,0.0003037559,0.35294807,0.0082562165,0.000031628402],"about_ca_topic_score_codex":0.015245871,"about_ca_topic_score_gemma":0.008045729,"teacher_disagreement_score":0.015245871,"about_ca_system_score_codex":0.003406578,"about_ca_system_score_gemma":0.005963183,"threshold_uncertainty_score":0.0615021},"labels":[],"label_agreement":null},{"id":"W4210294742","doi":"10.1109/ase51524.2021.9678520","title":"Subtle Bugs Everywhere: Generating Documentation for Data Wrangling Code","year":2021,"lang":"en","type":"article","venue":"2021 36th IEEE/ACM International Conference on Automated Software Engineering (ASE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Documentation; Code (set theory); Programming language; Reuse; Plug-in; Software bug; Database; Software engineering; Software; Engineering; Set (abstract data type)","score_opus":0.07020492043842771,"score_gpt":0.3442052189351064,"score_spread":0.27400029849667873,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210294742","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048732873,0.00022775849,0.7281107,0.00063254486,0.00017629714,0.0006041682,0.0027160503,0.21332934,0.0054702456],"genre_scores_gemma":[0.1759035,0.0002449866,0.7641879,0.0003582699,0.00006428248,0.0007065772,0.006920453,0.044447545,0.00716644],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964803,0.0008949452,0.00039457338,0.0006477146,0.001415122,0.0001673214],"domain_scores_gemma":[0.9633649,0.016951736,0.0028775204,0.011755317,0.0043978365,0.0006526492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047372524,0.0015558752,0.0005639002,0.002839777,0.00078943034,0.0018317517,0.002091566,0.0012972789,0.0070668063],"category_scores_gemma":[0.037257142,0.0013855625,0.0009914874,0.0013981222,0.0010102422,0.002981326,0.003048274,0.0019005528,0.0033989763],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008525086,0.00087594293,0.018171597,0.0016454142,0.00015110656,0.0021257238,0.0076438673,0.01778179,0.06844699,0.014070249,0.08685815,0.78137654],"study_design_scores_gemma":[0.0006822046,0.0013404515,0.011906907,0.001176003,0.000246298,0.0035500012,0.0015096895,0.39948985,0.29122156,0.03177271,0.2566788,0.00042551087],"about_ca_topic_score_codex":0.00090691994,"about_ca_topic_score_gemma":0.0015849681,"teacher_disagreement_score":0.0070668063,"about_ca_system_score_codex":0.00054004585,"about_ca_system_score_gemma":0.0015289882,"threshold_uncertainty_score":0.025053322},"labels":[],"label_agreement":null},{"id":"W4210729440","doi":"10.1016/j.infsof.2022.106855","title":"An empirical study on self-admitted technical debt in modern code review","year":2022,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Japan Society for the Promotion of Science; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung","keywords":"Technical debt; Debt; Imperfect; Code (set theory); Process (computing); Computer science; Empirical research; Risk analysis (engineering); Business; Accounting; Data science; Finance; Statistics; Mathematics; Software development","score_opus":0.017732388542163458,"score_gpt":0.31609050294286134,"score_spread":0.29835811440069787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210729440","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9968933,0.00033737646,0.00014411444,0.00027732967,0.0000067009373,0.000016249482,0.00011490584,0.000005002471,0.0022050405],"genre_scores_gemma":[0.9989767,0.00013910956,0.00007037712,0.00005879393,0.000015794472,0.000010403053,0.0001473862,0.000004663541,0.0005766879],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99633884,0.0015184131,0.00044226937,0.00032058638,0.0009407716,0.0004391083],"domain_scores_gemma":[0.7139327,0.1436362,0.103240214,0.006630568,0.021291986,0.011268236],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053281174,0.0001249199,0.00027755,0.0033349223,0.0014152398,0.0022995337,0.0007524782,0.0010065185,0.0033597925],"category_scores_gemma":[0.10803906,0.00024515085,0.00025086224,0.0041577816,0.0015417916,0.0027658632,0.001640123,0.0019297925,0.00043130948],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013454714,0.00025742583,0.98191065,0.000071581104,0.000040186307,0.00023261698,0.0069554313,0.00010639627,0.00018989266,0.001440761,0.00078711245,0.00787352],"study_design_scores_gemma":[0.000011224229,0.00012100487,0.98511666,0.00007619188,0.000040267947,0.00033984004,0.010064765,0.0005763589,0.00021351145,0.00045632347,0.0029641348,0.00001970558],"about_ca_topic_score_codex":0.013579218,"about_ca_topic_score_gemma":0.019952156,"teacher_disagreement_score":0.013579218,"about_ca_system_score_codex":0.0018114342,"about_ca_system_score_gemma":0.0023567586,"threshold_uncertainty_score":0.028178096},"labels":[],"label_agreement":null},{"id":"W4210732697","doi":"10.1016/j.jss.2022.111229","title":"Evaluating the performance of clone detection tools in detecting cloned co-change candidates","year":2022,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"clone (Java method); Java; Commit; Cloning (programming); Computer science; Computational biology; Software evolution; Source code; Biology; Software; Data mining; Genetics; Programming language; Software system; Gene; Database","score_opus":0.07439687648200605,"score_gpt":0.3285819411431132,"score_spread":0.25418506466110713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210732697","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9626777,0.0007887374,0.027615456,0.00009365912,0.00007854522,0.00008326641,0.00058187416,0.0069022346,0.0011785297],"genre_scores_gemma":[0.9422913,0.00019207729,0.054716192,0.0000567689,0.000025257848,0.00004079866,0.0014222515,0.00020991628,0.0010454437],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9947076,0.0011383139,0.0005764498,0.0011422876,0.0019536098,0.0004817732],"domain_scores_gemma":[0.9404847,0.04561977,0.0034053859,0.0028315133,0.0062422194,0.0014164636],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0044945786,0.0013645653,0.0010863148,0.0047279852,0.0006951669,0.001716208,0.0018728818,0.0023031489,0.0011140757],"category_scores_gemma":[0.03078885,0.0004341044,0.0008094523,0.0027823362,0.0006035267,0.0019112345,0.0009916637,0.0007997981,0.0007251557],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0082140025,0.0022139484,0.24543427,0.0013093194,0.0014899297,0.0013305142,0.00087469415,0.110501096,0.12315148,0.0016059362,0.0060384413,0.4978364],"study_design_scores_gemma":[0.0001885375,0.0019928575,0.041463323,0.000046213103,0.0003476142,0.00082283263,0.00037045818,0.8719621,0.08043998,0.00069586886,0.0015694768,0.00010076192],"about_ca_topic_score_codex":0.0064263456,"about_ca_topic_score_gemma":0.006586305,"teacher_disagreement_score":0.9955054,"about_ca_system_score_codex":0.00058060227,"about_ca_system_score_gemma":0.0012293807,"threshold_uncertainty_score":0.023769855},"labels":[],"label_agreement":null},{"id":"W4210736653","doi":"10.5753/wscad.2009.17394","title":"Automação de Refatorações para Programas Fortran de Alto Desempenho","year":2009,"lang":"pt","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ovarian Cancer Canada","funders":"","keywords":"Humanities; Physics; Computer science; Art","score_opus":0.038344187266792396,"score_gpt":0.33479600138263466,"score_spread":0.29645181411584226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210736653","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068441875,0.0011063892,0.8421917,0.0006095392,0.00021618309,0.0005057676,0.0007816945,0.06857552,0.017571358],"genre_scores_gemma":[0.3722934,0.00096120924,0.57420444,0.0005768485,0.000078935256,0.0005100221,0.0023491734,0.008975187,0.040050786],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99747974,0.00043394038,0.00022820865,0.00063532824,0.0009853618,0.00023740222],"domain_scores_gemma":[0.99477494,0.0016680869,0.0003615731,0.0017201063,0.0013338077,0.00014140959],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014631979,0.0014024436,0.00063007226,0.0010634217,0.001033396,0.002384195,0.0014702982,0.0009831411,0.011732017],"category_scores_gemma":[0.0076472955,0.0008615009,0.0013319503,0.00095580245,0.0009849094,0.0028952388,0.0019896813,0.0017279012,0.005309117],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008928227,0.00023135309,0.0058006262,0.0015014124,0.0001306958,0.0008444406,0.0038021777,0.018731214,0.112304464,0.021465216,0.016226033,0.81806946],"study_design_scores_gemma":[0.0002234953,0.0008416171,0.008341247,0.0005037716,0.0003139736,0.0017418186,0.0015792553,0.21843734,0.34697664,0.026197877,0.39453351,0.00030948845],"about_ca_topic_score_codex":0.00607902,"about_ca_topic_score_gemma":0.007990357,"teacher_disagreement_score":0.011732017,"about_ca_system_score_codex":0.0010122561,"about_ca_system_score_gemma":0.0019701603,"threshold_uncertainty_score":0.039247513},"labels":[],"label_agreement":null},{"id":"W4212906466","doi":"10.1109/tse.2022.3152148","title":"An Empirical Study of Yanked Releases in the Rust Package Registry","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); University of Alberta","funders":"University of Alberta","keywords":"Computer science; Rust (programming language); Software versioning; Dependency (UML); Software package; Code (set theory); Software release life cycle; Software; Software engineering; Operating system; World Wide Web; Database; Programming language; Software development; Software quality","score_opus":0.024242124774664125,"score_gpt":0.28821594235437353,"score_spread":0.2639738175797094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4212906466","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99674857,0.00017423982,0.0004091445,0.00024252894,0.0000072657062,0.00002270318,0.0002075879,0.000016116077,0.0021719784],"genre_scores_gemma":[0.9979538,0.00022268597,0.00045258683,0.00010124382,0.000010338625,0.000028965374,0.00041694293,0.000031171985,0.0007823065],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.991971,0.0032152534,0.00079500943,0.0010860183,0.002113384,0.0008192302],"domain_scores_gemma":[0.8389564,0.08298042,0.053733166,0.006711903,0.012932061,0.0046860725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010675109,0.00031199236,0.00032187812,0.0023763336,0.0012435244,0.0029841196,0.0011374695,0.0007748754,0.0034705915],"category_scores_gemma":[0.059205726,0.00045752173,0.00043391535,0.003779939,0.0022300347,0.007493709,0.0021044472,0.0026177072,0.0009904249],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015426023,0.00029116252,0.97039056,0.00012972484,0.000041971787,0.0006132322,0.012357914,0.00029988948,0.00040505722,0.0015234423,0.0014576991,0.012335141],"study_design_scores_gemma":[0.000017940656,0.0002278944,0.95045483,0.00015332489,0.000038083665,0.00067633274,0.03697499,0.0023937204,0.00052427134,0.0005498221,0.007936163,0.000052524992],"about_ca_topic_score_codex":0.0081381425,"about_ca_topic_score_gemma":0.009612053,"teacher_disagreement_score":0.010675109,"about_ca_system_score_codex":0.0016107208,"about_ca_system_score_gemma":0.0014575794,"threshold_uncertainty_score":0.05645603},"labels":[],"label_agreement":null},{"id":"W4213075374","doi":"10.1109/tse.2018.2872711","title":"An Interactive and Dynamic Search-Based Approach to Software Refactoring Recommendations","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Ford Motor Company","keywords":"Code refactoring; Computer science; Software; Software engineering; Software evolution; Software quality; Software system; Set (abstract data type); Process (computing); Merge (version control); Software metric; Software development; Programming language; Software construction; Information retrieval","score_opus":0.0253112702250691,"score_gpt":0.27451925036300423,"score_spread":0.24920798013793513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4213075374","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032135043,0.0010077901,0.9515459,0.0007706113,0.000100619305,0.00080377574,0.00036259668,0.00823389,0.00503971],"genre_scores_gemma":[0.2172736,0.00040615475,0.7751564,0.0005705136,0.00010554547,0.00089181354,0.0010267282,0.00037450928,0.0041948194],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964857,0.0012383852,0.00024322786,0.00073873595,0.0010675116,0.00022634993],"domain_scores_gemma":[0.9916998,0.005499951,0.0005089464,0.0006358892,0.0013524641,0.000303015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037344655,0.0026258468,0.002270312,0.004845952,0.0012129529,0.0016772674,0.0061908565,0.0038290918,0.005104146],"category_scores_gemma":[0.0116998,0.0012718793,0.0016227376,0.0027513104,0.00079587137,0.002129598,0.0021969227,0.0019883174,0.001728254],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081898854,0.0019836363,0.0073219878,0.00092755875,0.0006001513,0.000634213,0.0013508055,0.30628523,0.0225995,0.0072470196,0.017061409,0.6331694],"study_design_scores_gemma":[0.00015204285,0.00024034463,0.0008140076,0.000046521673,0.00013754409,0.00015655298,0.00013725007,0.9861704,0.002659017,0.004143494,0.005287973,0.00005479941],"about_ca_topic_score_codex":0.011082807,"about_ca_topic_score_gemma":0.028010545,"teacher_disagreement_score":0.011082807,"about_ca_system_score_codex":0.0011717823,"about_ca_system_score_gemma":0.0020108798,"threshold_uncertainty_score":0.022036552},"labels":[],"label_agreement":null},{"id":"W4213310871","doi":"10.7287/peerj.preprints.2723","title":"Lifting the curse of stringly-typed code","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"JavaScript; Computer science; String (physics); Programming language; Parsing; Code (set theory); Source code; Rule-based machine translation; Natural language processing; Theoretical computer science; Artificial intelligence; Mathematics; Set (abstract data type)","score_opus":0.05605265536245688,"score_gpt":0.3304924646234312,"score_spread":0.27443980926097433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4213310871","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9586183,0.0009404942,0.027704127,0.0017370892,0.00006573429,0.00008632653,0.0016727591,0.0022464893,0.006928701],"genre_scores_gemma":[0.9699874,0.00045829418,0.021710612,0.00084964774,0.000053858017,0.000058878275,0.0024756372,0.0018489611,0.0025565913],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.98039013,0.006260909,0.0013734268,0.0035102905,0.0076054744,0.0008597498],"domain_scores_gemma":[0.8350705,0.106501624,0.019798573,0.024967846,0.012165669,0.0014957815],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011164889,0.00063081883,0.0006067227,0.0043045394,0.0012644785,0.0030765936,0.0012765434,0.0015462649,0.0024676684],"category_scores_gemma":[0.1549184,0.0009372405,0.00068645464,0.0047887294,0.0030735487,0.008841553,0.0036436438,0.0028299557,0.0017339248],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009833084,0.0005094931,0.6109477,0.0013532913,0.00042973217,0.0028867421,0.03097145,0.0036816301,0.019473106,0.016736275,0.023620276,0.288407],"study_design_scores_gemma":[0.00015686611,0.000679389,0.5864114,0.0023293574,0.000629165,0.012389068,0.02957109,0.07965424,0.0460161,0.075490154,0.16629969,0.00037353413],"about_ca_topic_score_codex":0.0033913462,"about_ca_topic_score_gemma":0.005851195,"teacher_disagreement_score":0.011164889,"about_ca_system_score_codex":0.0008036347,"about_ca_system_score_gemma":0.0017931415,"threshold_uncertainty_score":0.05904627},"labels":[],"label_agreement":null},{"id":"W4214585125","doi":"10.1145/3511430.3511441","title":"Supporting Readability by Comprehending the Hierarchical Abstraction of a Software Project","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Readability; Computer science; Abstraction; Software engineering; Programming language; Software","score_opus":0.024094735349566892,"score_gpt":0.316441231488114,"score_spread":0.2923464961385471,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4214585125","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30686063,0.00072730385,0.66802335,0.0017662728,0.000078967336,0.00054102583,0.0012039634,0.005237657,0.015560694],"genre_scores_gemma":[0.52947927,0.0005858361,0.46359423,0.00030875488,0.00007140109,0.00040222972,0.0019247408,0.0008104083,0.0028231484],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971449,0.0012956646,0.00021317524,0.00050412875,0.00071402546,0.00012812068],"domain_scores_gemma":[0.97262436,0.018243024,0.0029359267,0.002654765,0.0031862392,0.00035569855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054508694,0.00085600873,0.0004643157,0.006161562,0.0009605886,0.0037477626,0.0009698619,0.0011481013,0.003110318],"category_scores_gemma":[0.04083777,0.00038174164,0.00064984284,0.0026174171,0.0011151673,0.010587743,0.002146798,0.0014904639,0.0013632163],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026927868,0.00022592886,0.028753694,0.0014478033,0.000050906983,0.0004575014,0.09682415,0.006344912,0.054460242,0.02220385,0.012179926,0.7767818],"study_design_scores_gemma":[0.00012345186,0.0014027504,0.13808075,0.0020574976,0.00032347106,0.0039024344,0.09433905,0.21417822,0.06955799,0.18131989,0.29401204,0.0007024623],"about_ca_topic_score_codex":0.0036405034,"about_ca_topic_score_gemma":0.004304353,"teacher_disagreement_score":0.006161562,"about_ca_system_score_codex":0.0010551417,"about_ca_system_score_gemma":0.0014043073,"threshold_uncertainty_score":0.02882731},"labels":[],"label_agreement":null},{"id":"W4214592544","doi":"10.1007/978-1-4842-8051-5_20","title":"How to Use the Recipes","year":2022,"lang":"en","type":"book-chapter","venue":"Apress eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Alberta Bible College","funders":"","keywords":"Recipe; Operationalization; Audit; Computer science; Code (set theory); Code of practice; Engineering; Programming language; Engineering management; Business; Accounting; History; Epistemology","score_opus":0.05233012107935242,"score_gpt":0.24696792805480633,"score_spread":0.1946378069754539,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4214592544","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011811487,0.0028058495,0.3651327,0.012439791,0.0043631806,0.0011448559,0.0014689491,0.020209465,0.591254],"genre_scores_gemma":[0.010416073,0.003915243,0.3718587,0.0049606906,0.00072262506,0.00076954765,0.0018867023,0.015102136,0.5903683],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9958769,0.000803787,0.00024700706,0.000430993,0.0024487737,0.00019247978],"domain_scores_gemma":[0.99571663,0.0013230188,0.00013110769,0.0009439106,0.0016587506,0.00022653012],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030422793,0.0019287906,0.0008271123,0.0021018982,0.0025671928,0.008963799,0.0023812617,0.0027304487,0.19027446],"category_scores_gemma":[0.016814884,0.0014187801,0.0009975405,0.001289103,0.0028183977,0.013055613,0.003644908,0.0065637766,0.21241416],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004876255,0.00011039386,0.00019771222,0.00050648896,0.00001547276,0.00021664712,0.0020377159,0.0013255686,0.0024090277,0.22407252,0.47008902,0.29897073],"study_design_scores_gemma":[0.0000040397954,0.0000103147195,0.00005013303,0.0001260896,0.0000028066845,0.00015981091,0.00019914034,0.00040101472,0.0008992975,0.021244334,0.97688055,0.000022583585],"about_ca_topic_score_codex":0.0030634315,"about_ca_topic_score_gemma":0.0040370845,"teacher_disagreement_score":0.19027446,"about_ca_system_score_codex":0.0015053443,"about_ca_system_score_gemma":0.0025739586,"threshold_uncertainty_score":0.6365315},"labels":[],"label_agreement":null},{"id":"W4214741785","doi":"10.1007/s10664-021-10070-w","title":"TraceSim: An Alignment Method for Computing Stack Trace Similarity","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Telefonaktiebolaget LM Ericsson; Western Canada Research Grid; Compute Canada","keywords":"Computer science; Data deduplication; Crash; Data mining; TRACE (psycholinguistics); Software; Task (project management); Subroutine; Flexibility (engineering); Benchmark (surveying); Similarity (geometry); Information retrieval; Artificial intelligence; Machine learning; Database; Programming language; Engineering","score_opus":0.04467345933549554,"score_gpt":0.34524277005109066,"score_spread":0.3005693107155951,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4214741785","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014188775,0.00024596552,0.9434972,0.0000660146,0.00017599195,0.00021084782,0.002769149,0.03731395,0.0015321407],"genre_scores_gemma":[0.09748919,0.00024032935,0.88075626,0.00008230707,0.00008475869,0.0006209624,0.0111682825,0.0059861634,0.0035717788],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9964282,0.0005203065,0.00046416646,0.0008275076,0.0014961144,0.00026371545],"domain_scores_gemma":[0.99485165,0.001499495,0.0005686981,0.0012430171,0.0015981591,0.00023897541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021726515,0.002268631,0.0016876182,0.009886033,0.0016234663,0.0030118262,0.003446664,0.0018681084,0.011372411],"category_scores_gemma":[0.018703416,0.0010830738,0.0016724197,0.010464676,0.00079746934,0.0052583176,0.0032683904,0.002578048,0.0058603836],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010529498,0.0005270039,0.01094595,0.0012482095,0.00072573277,0.00038105214,0.0012909959,0.023624498,0.034306686,0.022538656,0.04449526,0.85886306],"study_design_scores_gemma":[0.00030146693,0.0007383208,0.009259624,0.0002364883,0.00043878326,0.0011676252,0.0013693951,0.744424,0.08512905,0.0792137,0.07738351,0.0003380487],"about_ca_topic_score_codex":0.005397487,"about_ca_topic_score_gemma":0.009062468,"teacher_disagreement_score":0.011372411,"about_ca_system_score_codex":0.00083742384,"about_ca_system_score_gemma":0.002839408,"threshold_uncertainty_score":0.038044512},"labels":[],"label_agreement":null},{"id":"W4214871480","doi":"10.1007/s10664-021-10078-2","title":"Reuse and maintenance practices among divergent forks in three software ecosystems","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; Polytechnique Montréal","funders":"Vetenskapsrådet; Vlaamse regering; Fonds De La Recherche Scientifique - FNRS; Fonds Wetenschappelijk Onderzoek; Canada Research Chairs","keywords":"Computer science; Software engineering; Software development; Software construction; Software bug; Software; Social software engineering; Software evolution; Software analytics; Software peer review; Software system; Code reuse; Software maintenance; Operating system","score_opus":0.02947488547754301,"score_gpt":0.27531612184347026,"score_spread":0.24584123636592725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4214871480","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9951308,0.00014687103,0.0038002012,0.00004325217,0.0000013550925,0.000021866601,0.000095661344,0.00010041241,0.00065964734],"genre_scores_gemma":[0.9882599,0.000094754025,0.010579662,0.000018203391,0.000001686493,0.000026435562,0.00052276463,0.000040682837,0.0004559106],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99808407,0.00038756552,0.00018443137,0.00048007592,0.000611894,0.0002519671],"domain_scores_gemma":[0.9846704,0.007332593,0.0029535121,0.0016020385,0.0025848525,0.0008564931],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030378092,0.00028140153,0.00030859953,0.0041869413,0.001939734,0.0016516667,0.0007148413,0.0005890493,0.00077995],"category_scores_gemma":[0.01845285,0.00037946395,0.0006514535,0.0031035482,0.0015411298,0.0028908632,0.0020272501,0.00048010246,0.00017705448],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030377175,0.0001540282,0.8839244,0.00018528033,0.00011542627,0.0012966763,0.012832746,0.0046635033,0.0074269553,0.003683659,0.00093123456,0.084482305],"study_design_scores_gemma":[0.000053179414,0.00030180643,0.8782199,0.00022610964,0.00021181966,0.0028084745,0.018792707,0.064009465,0.008876556,0.0122969095,0.014070202,0.00013279135],"about_ca_topic_score_codex":0.009343545,"about_ca_topic_score_gemma":0.013858428,"teacher_disagreement_score":0.009343545,"about_ca_system_score_codex":0.0016396794,"about_ca_system_score_gemma":0.0013360921,"threshold_uncertainty_score":0.018578291},"labels":[],"label_agreement":null},{"id":"W4220704127","doi":"10.1145/3510455.3512771","title":"Better modeling the programming world with code concept graphs-augmented multi-modal learning","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Identifier; Artificial intelligence; Code (set theory); Domain (mathematical analysis); Process (computing); Machine learning; Software engineering; Programming language","score_opus":0.03777001590869307,"score_gpt":0.29157668098486095,"score_spread":0.2538066650761679,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220704127","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05529485,0.00018689866,0.9417173,0.0005310262,0.000024091925,0.000032955646,0.00015956484,0.0009344604,0.0011188278],"genre_scores_gemma":[0.7304604,0.00021009697,0.26653472,0.00025611994,0.000028176259,0.00009923152,0.0004863544,0.00014208052,0.0017828317],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99964666,0.0001294141,0.00001344255,0.00011564924,0.00006196223,0.0000328181],"domain_scores_gemma":[0.9984358,0.0009130023,0.00015451557,0.00024508176,0.00018206035,0.00006957253],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007354187,0.00063005456,0.00044723568,0.0008177867,0.00027958627,0.0011171199,0.0013652679,0.0010062376,0.0016010604],"category_scores_gemma":[0.0036570758,0.00036151215,0.00097385637,0.0007448544,0.00086957286,0.0036961967,0.0012583003,0.0024771956,0.00029992414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007466609,0.0001276028,0.0019229809,0.00008555041,0.00004376064,0.00007363072,0.00016382991,0.89615905,0.0036735453,0.026282752,0.0013013277,0.07009134],"study_design_scores_gemma":[0.0000019253469,0.0000091298425,0.00007395515,0.0000026583448,0.000002272537,0.0000064599662,0.000007801367,0.98791796,0.00031804087,0.011456933,0.0002002598,0.0000026191485],"about_ca_topic_score_codex":0.0066194898,"about_ca_topic_score_gemma":0.0104868105,"teacher_disagreement_score":0.0066194898,"about_ca_system_score_codex":0.0009583488,"about_ca_system_score_gemma":0.00082035875,"threshold_uncertainty_score":0.013161898},"labels":[],"label_agreement":null},{"id":"W4220908806","doi":"10.18280/ria.360109","title":"Novel Optimized Reusable Component Repository Using Neural Networks","year":2022,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Component (thermodynamics); Computer science; Software; Information repository; Artificial neural network; Work (physics); Component-based software engineering; Data science; World Wide Web; Software engineering; Database; Software development; Computer data storage; Engineering; Artificial intelligence; Operating system","score_opus":0.05333537265837602,"score_gpt":0.27729546277295297,"score_spread":0.22396009011457696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220908806","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09397286,0.0014374742,0.8907556,0.00033491192,0.00017811719,0.00015267388,0.00023104437,0.0060333526,0.0069039105],"genre_scores_gemma":[0.76369274,0.0007252447,0.22227602,0.0001567783,0.0000684821,0.0002535212,0.0006191243,0.0001637073,0.012044457],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99964285,0.000043780532,0.000029373216,0.00010500481,0.000102119586,0.00007692558],"domain_scores_gemma":[0.9995592,0.00011457782,0.0000752761,0.000057608344,0.00016594518,0.000027241233],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00059290486,0.00074542285,0.0011054587,0.00140263,0.0005609425,0.0011488671,0.0021180767,0.0009778744,0.002781642],"category_scores_gemma":[0.0012987414,0.0004797495,0.0010029896,0.0013859564,0.0002699706,0.0016779965,0.00084117707,0.000645049,0.0007759265],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029678788,0.00024190871,0.0017747072,0.00014806245,0.00013868186,0.0002240609,0.00006043235,0.5235929,0.008984061,0.0046633724,0.0047390675,0.455136],"study_design_scores_gemma":[0.000012025795,0.000046710386,0.00020304469,0.000008695524,0.000024304036,0.000043393327,0.000009602872,0.99584335,0.0021986056,0.0008688401,0.00073383364,0.000007537718],"about_ca_topic_score_codex":0.011546976,"about_ca_topic_score_gemma":0.00961524,"teacher_disagreement_score":0.011546976,"about_ca_system_score_codex":0.0010143844,"about_ca_system_score_gemma":0.001057889,"threshold_uncertainty_score":0.02295953},"labels":[],"label_agreement":null},{"id":"W4221004552","doi":"10.1142/s0218194022500085","title":"GASSER: A Multi-Objective Evolutionary Approach for Test Suite Reduction","year":2022,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Test suite; Sorting; Computer science; Regression testing; Reduction (mathematics); Suite; Genetic algorithm; Test case; Artifact (error); Baseline (sea); Software; Model-based testing; Code coverage; Machine learning; Evolutionary algorithm; Data mining; Artificial intelligence; Algorithm; Regression analysis; Software system; Mathematics; Programming language","score_opus":0.01419625295735528,"score_gpt":0.25685773653825233,"score_spread":0.24266148358089706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221004552","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03638527,0.0017744276,0.95382166,0.00034964085,0.00011567124,0.00039410687,0.00026305811,0.0029791477,0.003917019],"genre_scores_gemma":[0.23002689,0.00071092474,0.76388633,0.00036080967,0.0000544656,0.0008289719,0.00076384394,0.00038801128,0.0029797405],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99872595,0.00051778107,0.00006221601,0.00016687112,0.00043438433,0.00009270874],"domain_scores_gemma":[0.99799323,0.0014236211,0.00013318248,0.00011291661,0.00028215253,0.000054950608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019128589,0.00230428,0.0012512095,0.0023798551,0.00047959728,0.000769161,0.0023269188,0.0014490823,0.0022035073],"category_scores_gemma":[0.004340395,0.00073799735,0.001786829,0.0013525557,0.00072178355,0.0006806698,0.0010868085,0.0016549212,0.0003771542],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005516764,0.0001238244,0.0014375385,0.00018577627,0.00020547425,0.00012923554,0.00006835546,0.88576233,0.0034272617,0.0030567932,0.0015528437,0.10399534],"study_design_scores_gemma":[0.000039286988,0.0001018959,0.00038141297,0.000022893379,0.000046782057,0.0000448948,0.000020699059,0.99374735,0.0011021215,0.0025958128,0.001885575,0.000011364908],"about_ca_topic_score_codex":0.007159828,"about_ca_topic_score_gemma":0.0074405363,"teacher_disagreement_score":0.007159828,"about_ca_system_score_codex":0.00111389,"about_ca_system_score_gemma":0.0018155333,"threshold_uncertainty_score":0.014236331},"labels":[],"label_agreement":null},{"id":"W4221058968","doi":"10.32604/iasc.2022.027349","title":"Improve Representation for Cross-Language Clone Detection by Pretrain Using Tree Autoencoder","year":2022,"lang":"en","type":"article","venue":"Intelligent Automation & Soft Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Autoencoder; Artificial intelligence; Deep learning; Tree (set theory); Context (archaeology); Node (physics); Encoder; Embedding; Feature learning; Pattern recognition (psychology); Machine learning","score_opus":0.02670070496416517,"score_gpt":0.3361846129183926,"score_spread":0.3094839079542274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221058968","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12644614,0.00063903333,0.85609466,0.00031586978,0.00016686841,0.000114575465,0.00066041434,0.013016278,0.002546113],"genre_scores_gemma":[0.6690307,0.0004171997,0.314391,0.0007700983,0.00006279874,0.00028389916,0.0055966037,0.0008354652,0.008612285],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992046,0.00009086743,0.00005515232,0.0003430958,0.00019869203,0.00010764431],"domain_scores_gemma":[0.99861634,0.00039879803,0.00012809798,0.00026780277,0.0005338684,0.000055063156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010133345,0.0014875663,0.000882033,0.0010130791,0.0003826653,0.00079398404,0.001679387,0.00097120414,0.0016814652],"category_scores_gemma":[0.004075002,0.00050304324,0.0010193649,0.00081861863,0.0005095626,0.0025845254,0.0015154363,0.0019040058,0.0013139923],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024409642,0.00023815648,0.010265642,0.00015102357,0.00015653517,0.00034759927,0.0002767442,0.16410212,0.032157723,0.0030623456,0.009737155,0.7792608],"study_design_scores_gemma":[0.000011963117,0.00006705496,0.0012214796,0.000011785207,0.000027936465,0.00008293531,0.000036124075,0.9817202,0.013227878,0.0020513635,0.0015237388,0.000017499457],"about_ca_topic_score_codex":0.011444226,"about_ca_topic_score_gemma":0.014280825,"teacher_disagreement_score":0.011444226,"about_ca_system_score_codex":0.000827872,"about_ca_system_score_gemma":0.0014478416,"threshold_uncertainty_score":0.022755206},"labels":[],"label_agreement":null},{"id":"W4221069574","doi":"10.1016/j.jss.2022.111308","title":"Wayback Machine: A tool to capture the evolutionary behavior of the bug reports and their triage process in open-source software systems","year":2022,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Triage; Computer science; Process (computing); Software; Software engineering; Open source software; Software evolution; Open source; Operating system; Software system; Software construction; Medicine; Medical emergency","score_opus":0.014279740946357937,"score_gpt":0.25262178928616486,"score_spread":0.2383420483398069,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221069574","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15572223,0.00046620987,0.67033887,0.00024354829,0.00023960798,0.00039655587,0.007689794,0.16146325,0.0034399773],"genre_scores_gemma":[0.48512626,0.00029291213,0.49559042,0.00012715334,0.000076942415,0.0008001222,0.006930384,0.0070420327,0.0040139104],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99911076,0.00021044913,0.00008700845,0.00024319692,0.00027803352,0.00007053245],"domain_scores_gemma":[0.98760337,0.008716578,0.001381068,0.0014050406,0.00058637647,0.00030747833],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021186683,0.0010714582,0.0005695739,0.0038823646,0.00050360814,0.0013057988,0.0014906782,0.0013629224,0.004779802],"category_scores_gemma":[0.01643536,0.00072231266,0.000852828,0.002216453,0.000528289,0.002487522,0.0011151242,0.001279138,0.0010275812],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013183585,0.0008565262,0.12938303,0.0014807485,0.0008002625,0.0015918833,0.0050090444,0.14779414,0.03175701,0.03319562,0.049662784,0.5971505],"study_design_scores_gemma":[0.00006909916,0.00021255479,0.017181264,0.00006891956,0.00009417501,0.00029964943,0.00015499404,0.94221383,0.010862093,0.015564387,0.013191373,0.00008769926],"about_ca_topic_score_codex":0.004049987,"about_ca_topic_score_gemma":0.004978066,"teacher_disagreement_score":0.004779802,"about_ca_system_score_codex":0.0004092958,"about_ca_system_score_gemma":0.00088252855,"threshold_uncertainty_score":0.015990078},"labels":[],"label_agreement":null},{"id":"W4221130019","doi":"10.12688/openreseurope.14507.1","title":"An open-source natural language processing toolkit to support software development: addressing automatic bug detection, code summarisation and code search","year":2022,"lang":"en","type":"article","venue":"Open Research Europe","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Horizon 2020 Framework Programme; European Commission","keywords":"Computer science; Software engineering; Codebase; Code review; Source code; Parsing; Code generation; Documentation; Software; Code (set theory); KPI-driven code analysis; TRACE (psycholinguistics); Software development; Programming language; World Wide Web; Static program analysis; Operating system; Key (lock)","score_opus":0.09172202267189417,"score_gpt":0.38927321591392294,"score_spread":0.2975511932420288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221130019","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025518823,0.00058918906,0.51886195,0.0006074315,0.0002918424,0.0005520386,0.021930745,0.44717053,0.0074444236],"genre_scores_gemma":[0.02611975,0.0008505061,0.8032973,0.0006023149,0.00011710331,0.0010095884,0.090883285,0.06355022,0.013569945],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965875,0.0006803453,0.000513921,0.00069830887,0.0013390349,0.00018091546],"domain_scores_gemma":[0.99130696,0.004349689,0.00064357265,0.0014263298,0.0018412007,0.0004321727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002759198,0.002356467,0.0009933846,0.0038182575,0.0008704657,0.0034371892,0.0035662656,0.0022398746,0.030236656],"category_scores_gemma":[0.015953498,0.0015369969,0.002559121,0.0021151488,0.0011621264,0.0063341926,0.0051751765,0.003174801,0.030369585],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007339215,0.0003379129,0.0025630451,0.00715304,0.00032934372,0.0016701387,0.0024288674,0.009595003,0.03468822,0.031556383,0.4467173,0.46222678],"study_design_scores_gemma":[0.00039047198,0.00020976098,0.004012279,0.0011605703,0.00014829636,0.0019949302,0.0006218944,0.14111967,0.032788545,0.063652866,0.75345135,0.00044939658],"about_ca_topic_score_codex":0.0057360143,"about_ca_topic_score_gemma":0.011199582,"teacher_disagreement_score":0.030236656,"about_ca_system_score_codex":0.0011719638,"about_ca_system_score_gemma":0.0041395314,"threshold_uncertainty_score":0.101151705},"labels":[],"label_agreement":null},{"id":"W4223432363","doi":"10.1145/3524610.3529156","title":"Error identification strategies for Python Jupyter notebooks","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Debugging; Python (programming language); Computer science; Programming language; Exploratory data analysis; Software engineering; Data science; Data mining","score_opus":0.06175773032543677,"score_gpt":0.3416812496562637,"score_spread":0.27992351933082693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4223432363","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4513946,0.00037684737,0.51754546,0.0018829167,0.00013928596,0.0013710775,0.00043009155,0.018213063,0.008646619],"genre_scores_gemma":[0.6528185,0.00020922386,0.33563942,0.00083182525,0.0000371292,0.0010110104,0.0003760445,0.0024470678,0.0066298526],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96729213,0.018505387,0.0030323532,0.0047812597,0.0052574854,0.001131422],"domain_scores_gemma":[0.7080722,0.22771841,0.019461341,0.028059563,0.014203276,0.0024851828],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02598007,0.0019417123,0.00075929856,0.0023378145,0.001652852,0.004436462,0.0041202316,0.0024860543,0.0042613843],"category_scores_gemma":[0.20949958,0.0012047986,0.0007767218,0.0011443115,0.0030137342,0.008037114,0.0056972513,0.0024973704,0.0011334853],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029317322,0.0014533778,0.08589736,0.0026950683,0.00019996194,0.0063446434,0.25372422,0.01096567,0.04929097,0.03647914,0.017538233,0.5324796],"study_design_scores_gemma":[0.00082643394,0.004616853,0.08677582,0.005625156,0.0007049704,0.01743113,0.10799841,0.18963458,0.23567323,0.10993916,0.23888828,0.0018859548],"about_ca_topic_score_codex":0.0011795167,"about_ca_topic_score_gemma":0.0016474694,"teacher_disagreement_score":0.02598007,"about_ca_system_score_codex":0.0016081855,"about_ca_system_score_gemma":0.0026957789,"threshold_uncertainty_score":0.13739741},"labels":[],"label_agreement":null},{"id":"W4223464152","doi":"10.1016/j.infsof.2022.106912","title":"An empirical study of emoji use in software development communication","year":2022,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Instituto Nacional de Ciência e Tecnologia para Engenharia de Software; Fundação de Amparo à Ciência e Tecnologia do Estado de Pernambuco; Conselho Nacional de Desenvolvimento Científico e Tecnológico","keywords":"Emoji; Computer science; Software engineering; Software development; Empirical research; Software; World Wide Web; Programming language; Social media; Mathematics","score_opus":0.024901636920426556,"score_gpt":0.2952116746122025,"score_spread":0.2703100376917759,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4223464152","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99782676,0.000045205794,0.00015780255,0.00003328889,0.0000029847067,0.000010458663,0.0000109275,0.0000031646964,0.0019095093],"genre_scores_gemma":[0.9985812,0.00011167589,0.00047351533,0.00005725641,0.0000058142346,0.000034072935,0.000028854349,0.0000077787845,0.00069982925],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.995401,0.0029186246,0.00033723758,0.00023159482,0.00084503915,0.0002665513],"domain_scores_gemma":[0.85276955,0.12314408,0.012168269,0.0027381515,0.005899479,0.0032804105],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051698764,0.00030202544,0.00030901315,0.002120645,0.0018866827,0.0022836016,0.00062136765,0.0008582465,0.002153705],"category_scores_gemma":[0.05756722,0.00041509373,0.00022281031,0.0021235915,0.00096614106,0.0022853503,0.0017298588,0.0016579586,0.00039371086],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047311484,0.0026839168,0.7539805,0.00037302123,0.000085164975,0.00081338844,0.18806592,0.00014960209,0.004865483,0.001510824,0.0004552774,0.04654382],"study_design_scores_gemma":[0.000040207546,0.0013866683,0.84794015,0.00025345368,0.00014106836,0.0008419421,0.14110698,0.0016474697,0.0023233378,0.00034554486,0.0039228094,0.00005039733],"about_ca_topic_score_codex":0.0030284121,"about_ca_topic_score_gemma":0.005446952,"teacher_disagreement_score":0.0051698764,"about_ca_system_score_codex":0.0008232751,"about_ca_system_score_gemma":0.0011810631,"threshold_uncertainty_score":0.027341306},"labels":[],"label_agreement":null},{"id":"W4223561042","doi":"10.1145/3524610.3527920","title":"Backports","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Merge (version control); Open source; Source code; Software engineering; Code review; Open source software; Software; Software development; World Wide Web; Programming language; Software quality; Information retrieval","score_opus":0.027752340835281004,"score_gpt":0.2909284429458229,"score_spread":0.2631761021105419,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4223561042","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09523766,0.0016264687,0.043160763,0.0010201777,0.0006939222,0.0010002872,0.7091426,0.057984576,0.09013348],"genre_scores_gemma":[0.08821627,0.00085267785,0.04222871,0.000624366,0.00010862427,0.001251714,0.8209149,0.011229509,0.03457314],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9926737,0.0010847099,0.0010641806,0.0015604724,0.0030668736,0.0005499593],"domain_scores_gemma":[0.97756577,0.0045906543,0.0025855978,0.008985069,0.0055085295,0.0007644949],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031142994,0.0016631442,0.00078903243,0.0075977505,0.001554467,0.0037336284,0.001974321,0.0014369163,0.032695837],"category_scores_gemma":[0.023067141,0.0010296985,0.0013715586,0.0071279285,0.0006900454,0.0051667066,0.004169296,0.0019317203,0.036060132],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011085782,0.00037950306,0.06993747,0.0027978334,0.00024132496,0.0010223148,0.0025710403,0.003141184,0.005012597,0.018526547,0.5931668,0.30209476],"study_design_scores_gemma":[0.00006912765,0.00009509895,0.029331634,0.00038915063,0.000058478236,0.0005684055,0.0007834138,0.0024069021,0.0040239026,0.0053883055,0.95679724,0.00008829184],"about_ca_topic_score_codex":0.008788205,"about_ca_topic_score_gemma":0.011749511,"teacher_disagreement_score":0.032695837,"about_ca_system_score_codex":0.0009888249,"about_ca_system_score_gemma":0.0018923548,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4223598146","doi":"10.1145/3524610.3527886","title":"On the effectiveness of pretrained models for API learning","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; Kelowna General Hospital; University of British Columbia","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Natural language; Language model; Lexical analysis; Task (project management); Information retrieval; Context (archaeology); Encoder; Question answering; Transformer; Automatic summarization; Parsing; ENCODE; Programming language","score_opus":0.03599107396813926,"score_gpt":0.293690002250762,"score_spread":0.2576989282826227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4223598146","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3361261,0.024278577,0.5756951,0.009450257,0.0016583937,0.00051370036,0.0033171687,0.0154218245,0.033538982],"genre_scores_gemma":[0.8812171,0.0035215335,0.094805144,0.0022198658,0.00046185587,0.00036256204,0.0053702546,0.0008522285,0.011189456],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973736,0.0010527116,0.00019748954,0.0007846513,0.00035712784,0.00023441485],"domain_scores_gemma":[0.9746303,0.020515224,0.0005108496,0.002180802,0.0017542145,0.00040869234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007169214,0.00443526,0.0017075293,0.0014275475,0.0009874471,0.0024004106,0.0029534975,0.0041211084,0.006136402],"category_scores_gemma":[0.03627221,0.0011725158,0.0012422043,0.0011109948,0.0012109228,0.0074362983,0.0023814542,0.0059861294,0.0024298332],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014128153,0.0007654211,0.0067048017,0.00044174065,0.00049056066,0.00021460233,0.00012433647,0.6533466,0.002641661,0.009065745,0.017737528,0.30705422],"study_design_scores_gemma":[0.000035689165,0.00011169069,0.00037315683,0.000046103796,0.00005327766,0.00003155023,0.000024032888,0.99207616,0.0011991459,0.0054676305,0.0005681756,0.000013338432],"about_ca_topic_score_codex":0.022980107,"about_ca_topic_score_gemma":0.02081067,"teacher_disagreement_score":0.022980107,"about_ca_system_score_codex":0.0021702545,"about_ca_system_score_gemma":0.0023593365,"threshold_uncertainty_score":0.045692682},"labels":[],"label_agreement":null},{"id":"W4223971160","doi":"10.1016/j.infsof.2022.107025","title":"S-DABT: Schedule and Dependency-aware Bug Triage in open-source bug tracking systems","year":2022,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software regression; Software bug; Schedule; Dependency (UML); Dependency graph; Scheduling (production processes); Open source; Software; Software development; Software quality; Software engineering; Programming language; Operating system; Engineering","score_opus":0.014490333728166742,"score_gpt":0.25314797987584403,"score_spread":0.2386576461476773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4223971160","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09162066,0.0011521012,0.42810783,0.0008295109,0.0005903356,0.0005192048,0.0035781923,0.46780947,0.0057926127],"genre_scores_gemma":[0.5125187,0.0004898349,0.45651287,0.00040878716,0.0001799505,0.00035747673,0.009249354,0.014050803,0.006232264],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99650204,0.00086861296,0.00033410304,0.00080675364,0.0011779161,0.000310605],"domain_scores_gemma":[0.98817164,0.004402991,0.0013049109,0.0039040567,0.0014143372,0.0008021287],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00340634,0.0017553354,0.0010800415,0.0038900196,0.0010167515,0.0021411087,0.003150921,0.0013023727,0.005419084],"category_scores_gemma":[0.018389441,0.0011554902,0.0015153186,0.0018865984,0.000978531,0.003638311,0.003926804,0.0019722634,0.0023459485],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002154269,0.00085799256,0.024035811,0.0011363595,0.00039919576,0.0005642608,0.0012413305,0.040419962,0.030806843,0.008479817,0.104214184,0.78569],"study_design_scores_gemma":[0.0009960152,0.0010973008,0.01134718,0.00018253097,0.0003170387,0.0006524196,0.00037335244,0.8808771,0.04246896,0.015515638,0.045918405,0.00025414888],"about_ca_topic_score_codex":0.008571441,"about_ca_topic_score_gemma":0.012104347,"teacher_disagreement_score":0.008571441,"about_ca_system_score_codex":0.0010120758,"about_ca_system_score_gemma":0.0029269266,"threshold_uncertainty_score":0.018128633},"labels":[],"label_agreement":null},{"id":"W4224287853","doi":"10.1145/3508479","title":"Just-In-Time Defect Prediction on JavaScript Projects: A Replication Study","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Natural Science Foundation of China; National Research Foundation Singapore","keywords":"JavaScript; Computer science; Java; Unobtrusive JavaScript; Machine learning; Artificial intelligence; Replication (statistics); Java Programming Language; Natural language processing; Software engineering; Programming language; Rich Internet application","score_opus":0.10556455715455794,"score_gpt":0.3332004721546199,"score_spread":0.22763591500006197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224287853","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.983952,0.00093919877,0.009098691,0.0002717773,0.00011020262,0.00028127033,0.0030414695,0.0010666032,0.0012387136],"genre_scores_gemma":[0.97171164,0.00041198585,0.015022888,0.0002163601,0.000098673314,0.00031937062,0.010810049,0.00016401455,0.0012450606],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9923062,0.002552067,0.0006836667,0.0023124795,0.0017678206,0.00037783667],"domain_scores_gemma":[0.9323357,0.026679413,0.007964983,0.015758444,0.015398904,0.0018627098],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008426437,0.0011905794,0.001016876,0.0036038447,0.0007784706,0.0014538171,0.0018141805,0.0011895417,0.0008220244],"category_scores_gemma":[0.03992362,0.000445402,0.0015137431,0.002513006,0.00078963867,0.003664635,0.0012365668,0.0021052845,0.0009972227],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001134696,0.003721781,0.7818263,0.001358107,0.0008315272,0.00096700585,0.0025191668,0.014382587,0.0065614786,0.0009155834,0.016441718,0.16934006],"study_design_scores_gemma":[0.00033717928,0.0034063014,0.7386821,0.0004982477,0.000870725,0.0018859748,0.0040490525,0.21871385,0.010007953,0.0021932651,0.019085465,0.00026996373],"about_ca_topic_score_codex":0.009919353,"about_ca_topic_score_gemma":0.010368378,"teacher_disagreement_score":0.9915736,"about_ca_system_score_codex":0.0009346705,"about_ca_system_score_gemma":0.0013236507,"threshold_uncertainty_score":0.04456383},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"simulation_or_modeling","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"},{"model":"gpt","categories":[],"domain":null,"study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"split"},{"id":"W4224436806","doi":"10.1007/s10664-022-10133-6","title":"Revisiting reopened bugs in open source software systems","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Queen's University","funders":"","keywords":"Software bug; Rework; Documentation; Pipeline (software); Software; Open source; Computer science; Software engineering; Engineering; Computer security; Operating system; Embedded system","score_opus":0.03242953496458023,"score_gpt":0.29263426994624303,"score_spread":0.2602047349816628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224436806","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9741431,0.0015901771,0.012009648,0.0025499554,0.000165854,0.000056930872,0.00026706932,0.00021745198,0.008999792],"genre_scores_gemma":[0.992845,0.00038086317,0.0047133244,0.0002024258,0.000060458347,0.000015968708,0.00024273543,0.0001510893,0.0013880816],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99200296,0.0022353516,0.0007930174,0.0011277463,0.0033580633,0.00048282839],"domain_scores_gemma":[0.68808305,0.21334286,0.03897568,0.022586934,0.034412015,0.002599438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009264235,0.0004147838,0.00039511028,0.0055109905,0.0013508655,0.003540713,0.0015731854,0.001430043,0.0048830113],"category_scores_gemma":[0.23523413,0.000434174,0.0005092422,0.0034405359,0.0030091729,0.00903983,0.0027172035,0.002719958,0.00043021125],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006459558,0.00095324515,0.5711127,0.0014119537,0.0002521319,0.0020487183,0.03538069,0.004958433,0.007682009,0.04702421,0.007210461,0.32131952],"study_design_scores_gemma":[0.00011401839,0.0009318228,0.77583593,0.0023781827,0.0004984571,0.0024397091,0.036117755,0.038345963,0.010633381,0.095438175,0.03705132,0.00021527392],"about_ca_topic_score_codex":0.010783494,"about_ca_topic_score_gemma":0.016110009,"teacher_disagreement_score":0.010783494,"about_ca_system_score_codex":0.0020275523,"about_ca_system_score_gemma":0.0028598513,"threshold_uncertainty_score":0.04899454},"labels":[],"label_agreement":null},{"id":"W4225094266","doi":"10.1145/3491102.3501895","title":"Supercharging Trial-and-Error for Learning Complex Software Applications","year":2022,"lang":"en","type":"article","venue":"CHI Conference on Human Factors in Computing Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Autodesk (Canada); University of Waterloo","funders":"","keywords":"Computer science; Software; Proof of concept; Key (lock); Software engineering; Machine learning; Data science; Programming language; Computer security","score_opus":0.1725917647466699,"score_gpt":0.3652937626882987,"score_spread":0.19270199794162882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225094266","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0212482,0.00033329733,0.96388495,0.0006064978,0.00010190949,0.000449438,0.00004205074,0.010247538,0.0030860258],"genre_scores_gemma":[0.1812861,0.00029567446,0.81191206,0.00029147032,0.000055760232,0.001054312,0.00010853439,0.0015597653,0.003436311],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97399014,0.015301073,0.001407218,0.0033093905,0.005061028,0.00093105243],"domain_scores_gemma":[0.8517625,0.10805055,0.0060913707,0.024583144,0.007471116,0.0020413327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026685864,0.0021089802,0.0014623175,0.0022283006,0.0016403713,0.004467254,0.00784261,0.0026917667,0.010053939],"category_scores_gemma":[0.15082434,0.0014966498,0.0011191192,0.0012245727,0.0071152137,0.01252253,0.009434592,0.005010175,0.0029347034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012121748,0.0009527134,0.0050537684,0.001095702,0.0001276287,0.0003879412,0.008303394,0.043763727,0.016250074,0.07810511,0.009738579,0.83500916],"study_design_scores_gemma":[0.00075325306,0.0022392534,0.0018806461,0.000980192,0.00017068713,0.00096330815,0.0018059188,0.5371883,0.06978351,0.3092409,0.074504435,0.0004895851],"about_ca_topic_score_codex":0.0012297385,"about_ca_topic_score_gemma":0.0018447654,"teacher_disagreement_score":0.026685864,"about_ca_system_score_codex":0.0014471778,"about_ca_system_score_gemma":0.00426681,"threshold_uncertainty_score":0.14113003},"labels":[],"label_agreement":null},{"id":"W4225163285","doi":"10.1145/3522587","title":"There’s no Such Thing as a Free Lunch: Lessons Learned from Exploring the Overhead Introduced by the Greenkeeper Dependency Bot in Npm","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Queen's University","funders":"","keywords":"Computer science; Dependency (UML); Overhead (engineering); Process (computing); Computer security; Pipeline (software); Action (physics); Software; Software engineering; Risk analysis (engineering); Process management","score_opus":0.17598726110486262,"score_gpt":0.33581605162678857,"score_spread":0.15982879052192595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225163285","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.911591,0.0015797759,0.06517582,0.008626724,0.000091552356,0.00017295126,0.0003104874,0.0012726876,0.011179024],"genre_scores_gemma":[0.9670911,0.0004262068,0.030011233,0.0006114077,0.000033751006,0.00006917422,0.0002398861,0.00032911816,0.0011881243],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99277675,0.0038164507,0.00022346804,0.0009927703,0.0016003802,0.0005901828],"domain_scores_gemma":[0.9248485,0.058243386,0.0048120446,0.005973252,0.0046509397,0.0014719627],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010034121,0.00070069043,0.0006704211,0.0018619364,0.0016492617,0.0032756173,0.002032736,0.0014213563,0.0014573794],"category_scores_gemma":[0.06302603,0.0006346079,0.00042537195,0.0013360521,0.002999458,0.009369755,0.0023772244,0.002755816,0.0004077997],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011589151,0.0020397657,0.44316465,0.0018383249,0.00024426,0.00444571,0.064522915,0.02900301,0.014792513,0.037343502,0.018041825,0.38340467],"study_design_scores_gemma":[0.00018698088,0.0015201119,0.286957,0.0019015085,0.00029267187,0.0036272937,0.09060454,0.44250762,0.012490805,0.10750028,0.05201939,0.0003918338],"about_ca_topic_score_codex":0.0097435005,"about_ca_topic_score_gemma":0.01735904,"teacher_disagreement_score":0.010034121,"about_ca_system_score_codex":0.0024891603,"about_ca_system_score_gemma":0.0022131347,"threshold_uncertainty_score":0.053066075},"labels":[],"label_agreement":null},{"id":"W4225333936","doi":"10.5267/j.ijdns.2022.2.011","title":"Understanding and predicting bugs fixed by API-migrations","year":2022,"lang":"en","type":"article","venue":"International Journal of Data and Network Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Software bug; Computer science; Context (archaeology); Security bug; Schedule; Software engineering; Software; Computer security; Programming language; Biology; Operating system","score_opus":0.08412188676902072,"score_gpt":0.3226650559591267,"score_spread":0.23854316919010599,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225333936","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9885095,0.00060917164,0.008472906,0.00018685602,0.00002418545,0.0000881434,0.00064590375,0.0005012114,0.00096218125],"genre_scores_gemma":[0.9841002,0.0003446025,0.013490017,0.000031864216,0.0000206896,0.00005637076,0.0012801944,0.00008852404,0.0005876372],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9934941,0.0013279117,0.0010397483,0.0013118316,0.0024119283,0.00041454978],"domain_scores_gemma":[0.841075,0.08413737,0.050959118,0.0058699185,0.015695794,0.0022627476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006937134,0.0009243314,0.0004714855,0.0095069595,0.00054697826,0.0022687356,0.0009611835,0.0009885281,0.00091726583],"category_scores_gemma":[0.09647092,0.00051116507,0.00061071396,0.0034233516,0.00060584757,0.0032683858,0.0012438962,0.00093281636,0.00038935934],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009871125,0.00012303186,0.9368909,0.00020735263,0.000059297603,0.00042468408,0.002342726,0.001943577,0.0018632875,0.0002083173,0.0007876101,0.05505049],"study_design_scores_gemma":[0.000017327344,0.00051220995,0.94781786,0.00018216217,0.0001889114,0.001293787,0.003966641,0.038998604,0.003160408,0.00066716975,0.003119163,0.00007573818],"about_ca_topic_score_codex":0.010017911,"about_ca_topic_score_gemma":0.013912606,"teacher_disagreement_score":0.010017911,"about_ca_system_score_codex":0.00078292197,"about_ca_system_score_gemma":0.0010694243,"threshold_uncertainty_score":0.036687493},"labels":[],"label_agreement":null},{"id":"W4225383538","doi":"10.32473/flairs.v35i.130643","title":"Learning to Rank with BERT for Argument Quality Evaluation","year":2022,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Argument (complex analysis); Leverage (statistics); Ranking (information retrieval); Rank (graph theory); Computer science; Learning to rank; Pairwise comparison; Artificial intelligence; Quality (philosophy); Machine learning; Representation (politics); Task (project management); Mathematics; Epistemology; Political science; Engineering","score_opus":0.17335816357690312,"score_gpt":0.4124483342358819,"score_spread":0.23909017065897878,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225383538","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.104975104,0.015491535,0.7627727,0.0032991236,0.0017923225,0.0012302368,0.012601072,0.06309288,0.03474503],"genre_scores_gemma":[0.53831786,0.0014507596,0.41779152,0.0010142946,0.0008840281,0.0007173758,0.023528231,0.0020586934,0.014237221],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9878461,0.0058606523,0.0008078465,0.0014016221,0.0034099277,0.00067390397],"domain_scores_gemma":[0.96726704,0.021514148,0.0022722397,0.0041569467,0.0038164877,0.0009731176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013313824,0.0040838593,0.0025325117,0.00779169,0.0012963908,0.004887446,0.003469139,0.0055609737,0.01239752],"category_scores_gemma":[0.04980415,0.0006579913,0.001715988,0.0037087626,0.0015704877,0.00578361,0.0024028472,0.004842697,0.009445498],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015995004,0.0008104374,0.011712115,0.001701674,0.0004820791,0.0002857404,0.00034854334,0.18131267,0.0048789154,0.022260867,0.11592018,0.6586873],"study_design_scores_gemma":[0.00017294765,0.00038270457,0.0019261267,0.0001572175,0.000065072396,0.00018683112,0.0001399243,0.96004564,0.0047213132,0.019939922,0.012187825,0.00007443786],"about_ca_topic_score_codex":0.0039339643,"about_ca_topic_score_gemma":0.010220524,"teacher_disagreement_score":0.013313824,"about_ca_system_score_codex":0.002596239,"about_ca_system_score_gemma":0.0024995212,"threshold_uncertainty_score":0.070411086},"labels":[],"label_agreement":null},{"id":"W4225630774","doi":"10.48550/arxiv.2201.02215","title":"On the Prevalence, Impact, and Evolution of SQL Code Smells in Data-Intensive Systems","year":2022,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Wetenschappelijk Onderzoek; Fonds De La Recherche Scientifique - FNRS","keywords":"Code smell; Computer science; Code refactoring; SQL; Code (set theory); Code review; Software quality; Software; Software engineering; Database; Programming language; Software development","score_opus":0.09521332136890527,"score_gpt":0.2315162346778636,"score_spread":0.13630291330895833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225630774","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99766266,0.0003036856,0.0011521549,0.00013986991,0.0000040411414,0.000013962969,0.00016216749,0.000036040983,0.0005254044],"genre_scores_gemma":[0.9979493,0.00018004468,0.001245203,0.000036802725,0.0000088835595,0.000014850103,0.00034254385,0.000018263865,0.00020412791],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9919648,0.001935899,0.0010762063,0.0015976187,0.0028671287,0.0005583232],"domain_scores_gemma":[0.7737853,0.12393476,0.07305929,0.0071353125,0.018295225,0.0037900854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00735599,0.00037292717,0.00029294996,0.0048752837,0.00057996676,0.0016158643,0.0005962055,0.00082514784,0.0010125397],"category_scores_gemma":[0.07364759,0.00039279426,0.0005056349,0.003397334,0.0011681608,0.0032799684,0.0017066177,0.0010422906,0.00031561227],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000059482903,0.00006687052,0.9843487,0.00007416136,0.000047329882,0.00012804601,0.0016999212,0.00039224364,0.00080291356,0.00012237816,0.00012859634,0.012129307],"study_design_scores_gemma":[0.000002401707,0.000097287746,0.99553853,0.000039064744,0.000026504846,0.00021344017,0.0011456211,0.0019421926,0.0004914722,0.00018178899,0.00030702134,0.000014614811],"about_ca_topic_score_codex":0.003812052,"about_ca_topic_score_gemma":0.005778937,"teacher_disagreement_score":0.00735599,"about_ca_system_score_codex":0.00074180734,"about_ca_system_score_gemma":0.0007642399,"threshold_uncertainty_score":0.03890264},"labels":[],"label_agreement":null},{"id":"W4225714444","doi":"10.1007/s10664-022-10125-6","title":"Tracking bad updates in mobile apps: a search-based approach","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University; Queen's University; École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Android (operating system); Computer science; Benchmark (surveying); Genetic programming; Machine learning; Sorting; Mobile apps; Artificial intelligence; World Wide Web","score_opus":0.0278380218906817,"score_gpt":0.2811678476185983,"score_spread":0.2533298257279166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225714444","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7395105,0.013187367,0.20575517,0.0031977985,0.00046554412,0.001448154,0.011468517,0.0049152435,0.02005182],"genre_scores_gemma":[0.9396834,0.0011799585,0.04982265,0.00037548953,0.00023035625,0.00017012689,0.0034920042,0.00013056304,0.0049154647],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9928349,0.0013120369,0.0007043661,0.0015074832,0.0031620692,0.00047909844],"domain_scores_gemma":[0.96330196,0.024020134,0.004061481,0.0026193971,0.0052065547,0.0007904511],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041549,0.0014974352,0.0029452988,0.017952418,0.001364293,0.004374091,0.0032952449,0.004001861,0.0031607132],"category_scores_gemma":[0.035126135,0.0007482301,0.0014588556,0.010671931,0.0009336322,0.0067162104,0.0025215358,0.0016855194,0.0018730769],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022371826,0.0025184357,0.35848644,0.0020611829,0.0012866027,0.0014559028,0.0018356026,0.029350925,0.011510709,0.0103651285,0.020777637,0.55811423],"study_design_scores_gemma":[0.00017648835,0.0016616689,0.14718822,0.00042277836,0.0014617081,0.003649278,0.0029411863,0.79243416,0.012623606,0.02668721,0.01045599,0.0002977666],"about_ca_topic_score_codex":0.012468879,"about_ca_topic_score_gemma":0.02214521,"teacher_disagreement_score":0.017952418,"about_ca_system_score_codex":0.0010579374,"about_ca_system_score_gemma":0.0017120145,"threshold_uncertainty_score":0.024792612},"labels":[],"label_agreement":null},{"id":"W4225807728","doi":"10.1109/qrs54544.2021.00063","title":"Predictors of Software Metric Correlation: A Non-parametric Analysis","year":2021,"lang":"en","type":"article","venue":"2021 IEEE 21st International Conference on Software Quality, Reliability and Security (QRS)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Cyclomatic complexity; Metric (unit); Code (set theory); Correlation; Computer science; Code refactoring; Source lines of code; Code coverage; Software quality; Software metric; Software; Statistics; Programming language; Mathematics; Software development; Set (abstract data type); Engineering","score_opus":0.034565574273561835,"score_gpt":0.32378986542613564,"score_spread":0.2892242911525738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225807728","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9708711,0.00040497154,0.023862585,0.0003679276,0.00010262306,0.00022189741,0.0010268439,0.0003278496,0.0028141926],"genre_scores_gemma":[0.9957419,0.000045818164,0.0027718102,0.000036910566,0.000030952957,0.00034200348,0.0005130956,0.000070847236,0.0004465757],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.96325815,0.023326924,0.001953811,0.0044629104,0.005112452,0.0018857365],"domain_scores_gemma":[0.56594366,0.3984325,0.012228957,0.013988676,0.0065383115,0.0028678516],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028606111,0.0011577124,0.0014141579,0.0038670443,0.0008249706,0.0023556869,0.0022443207,0.0017548358,0.009391672],"category_scores_gemma":[0.1380795,0.00048180055,0.0027235511,0.0036885836,0.002806222,0.0037436222,0.0032414785,0.003846492,0.0018475701],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026505394,0.0008028257,0.93505883,0.0003133212,0.002298481,0.0005783679,0.0038017374,0.0050962428,0.0018500024,0.0025388203,0.0028767013,0.042134315],"study_design_scores_gemma":[0.00016534881,0.005200429,0.86371,0.0001565176,0.00074310025,0.0015281994,0.0049908664,0.11031739,0.0020724004,0.006896439,0.0039717373,0.00024752226],"about_ca_topic_score_codex":0.0010942208,"about_ca_topic_score_gemma":0.00046729392,"teacher_disagreement_score":0.028606111,"about_ca_system_score_codex":0.00053679675,"about_ca_system_score_gemma":0.0009701438,"threshold_uncertainty_score":0.15128541},"labels":[],"label_agreement":null},{"id":"W4225845650","doi":"10.2139/ssrn.4070797","title":"What are the Characteristics of Highly-Selected Packages? A Case Study on the NPM Ecosystem","year":2022,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Ecosystem; Environmental science; Environmental resource management; Ecology; Biology","score_opus":0.014763968480411864,"score_gpt":0.24918403620970775,"score_spread":0.23442006772929588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225845650","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98943055,0.000072586394,0.003118492,0.00058452267,0.0000051530674,0.00004380069,0.00005446911,0.000030823307,0.006659521],"genre_scores_gemma":[0.9938391,0.00007592013,0.003825442,0.000091971364,0.0000057183893,0.000022890688,0.00006304189,0.000028137452,0.0020478829],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99799335,0.00092270283,0.00006018459,0.00018042803,0.00050103775,0.00034238116],"domain_scores_gemma":[0.993604,0.0036513084,0.0008047627,0.00042885286,0.0008649069,0.00064627355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029055772,0.00023509591,0.00028519472,0.0012428298,0.0022603613,0.0024767537,0.0009393776,0.0016212951,0.0024084912],"category_scores_gemma":[0.011098322,0.00019046216,0.00033429707,0.0024046768,0.0013634819,0.002912001,0.0019059633,0.00067048403,0.00036410848],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006746965,0.0016939391,0.5701786,0.00045349888,0.00007440048,0.03655241,0.078927726,0.006037182,0.013265801,0.029451152,0.00697994,0.25571066],"study_design_scores_gemma":[0.00009553954,0.001225138,0.62790304,0.0004838927,0.00017188641,0.024725048,0.18535884,0.04657699,0.0077419514,0.023504855,0.08203209,0.00018065858],"about_ca_topic_score_codex":0.00750496,"about_ca_topic_score_gemma":0.01674534,"teacher_disagreement_score":0.00750496,"about_ca_system_score_codex":0.0017775631,"about_ca_system_score_gemma":0.0016023215,"threshold_uncertainty_score":0.015366375},"labels":[],"label_agreement":null},{"id":"W4225878264","doi":"10.1109/tse.2022.3162985","title":"Static Profiling of Alloy Models","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Correctness; Modeling language; Natural language processing; Profiling (computer programming); Matching (statistics); Programming language; Software; Artificial intelligence; Data science; Software engineering","score_opus":0.02111579023518419,"score_gpt":0.23700404072038123,"score_spread":0.21588825048519703,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225878264","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4088794,0.0012068754,0.5074705,0.00089024764,0.0001885393,0.00047155778,0.008408655,0.022092577,0.05039166],"genre_scores_gemma":[0.7159907,0.0007045236,0.24422202,0.00023933478,0.000057883895,0.00042207062,0.01848922,0.005217573,0.014656664],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9956151,0.0011145442,0.0003516233,0.0005931279,0.0020762787,0.00024923947],"domain_scores_gemma":[0.99038386,0.0042041535,0.000620357,0.0023121359,0.002358592,0.00012100979],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002538957,0.0007842832,0.0007539708,0.0034086152,0.0011508228,0.003000672,0.0013590934,0.00088660617,0.0057470277],"category_scores_gemma":[0.018309066,0.0008848462,0.0012605661,0.002954527,0.0007066035,0.0036012856,0.0015728478,0.0008747418,0.002338997],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010946877,0.00043049792,0.06636906,0.0013870876,0.0002715959,0.0016561463,0.011680538,0.200815,0.056666806,0.21577603,0.031670127,0.41218242],"study_design_scores_gemma":[0.000041414754,0.00017260328,0.010080153,0.00017276945,0.00014973554,0.00060785626,0.0014854175,0.7923206,0.049567506,0.035665393,0.109631285,0.000105273575],"about_ca_topic_score_codex":0.008989217,"about_ca_topic_score_gemma":0.014743637,"teacher_disagreement_score":0.008989217,"about_ca_system_score_codex":0.0021015161,"about_ca_system_score_gemma":0.0019701284,"threshold_uncertainty_score":0.019225717},"labels":[],"label_agreement":null},{"id":"W4225908028","doi":"10.1145/3524842.3527957","title":"How heated is it?","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Polytechnique Montréal","funders":"","keywords":"Oracle; Lock (firearm); Computer science; Variety (cybernetics); Code (set theory); Open source; Software; Software engineering; Engineering; Artificial intelligence","score_opus":0.04831979751246629,"score_gpt":0.3038914363655416,"score_spread":0.2555716388530753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225908028","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63093656,0.00880503,0.05782376,0.11062175,0.0037682715,0.000349432,0.00071538764,0.0014326642,0.1855471],"genre_scores_gemma":[0.9693132,0.0014768853,0.005641836,0.008231538,0.00036717206,0.00013821894,0.00022979298,0.00061592896,0.013985356],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.9765827,0.013046439,0.00082710333,0.0031310588,0.004408644,0.0020040686],"domain_scores_gemma":[0.9532641,0.025751682,0.0069347164,0.0037125908,0.0064032665,0.0039335717],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015932605,0.0008139488,0.0005818243,0.0023845262,0.009789098,0.013776474,0.0015577443,0.0032197025,0.009738587],"category_scores_gemma":[0.07209442,0.0007648906,0.0007456014,0.0020615268,0.013056535,0.019748775,0.009196052,0.006000708,0.002636995],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002585148,0.00007613993,0.049631637,0.0009929445,0.000106876774,0.0015613635,0.7347251,0.00024547632,0.005520824,0.06716465,0.0328611,0.10685538],"study_design_scores_gemma":[0.000016583424,0.000078751815,0.024858413,0.0012782956,0.00008002158,0.0012405354,0.63830423,0.0008716828,0.0022448418,0.05956143,0.27129525,0.00017001812],"about_ca_topic_score_codex":0.0032965997,"about_ca_topic_score_gemma":0.0035412044,"teacher_disagreement_score":0.015932605,"about_ca_system_score_codex":0.004100188,"about_ca_system_score_gemma":0.00348642,"threshold_uncertainty_score":0.0842607},"labels":[],"label_agreement":null},{"id":"W4226069469","doi":"10.1109/qrs54544.2021.00012","title":"Analyzing Structural Security Posture to Evaluate System Design Decisions","year":2021,"lang":"en","type":"article","venue":"2021 IEEE 21st International Conference on Software Quality, Reliability and Security (QRS)","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Leverage (statistics); Computer science; Computer security; Software security assurance; Security testing; Computer security model; Identification (biology); Resource (disambiguation); Security information and event management; Secure coding; Security service; Security through obscurity; Software; Cloud computing security; Risk analysis (engineering); Information security; Artificial intelligence; Cloud computing; Business","score_opus":0.07041286786716698,"score_gpt":0.3641071605009214,"score_spread":0.2936942926337544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226069469","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63999563,0.00037602105,0.33983865,0.00064680073,0.000059431655,0.00067786244,0.00082958053,0.00200625,0.015569696],"genre_scores_gemma":[0.8468289,0.000097608914,0.15160152,0.000049576964,0.0000128003285,0.00025348656,0.0004908075,0.00008861765,0.000576576],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99043244,0.003776015,0.00093607395,0.000621371,0.0037142453,0.00051990617],"domain_scores_gemma":[0.95486385,0.020681422,0.009416785,0.003744851,0.010370939,0.0009221224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010902034,0.0016185767,0.00087663805,0.009062931,0.0008065822,0.0031882247,0.00071825157,0.0011800252,0.0019173758],"category_scores_gemma":[0.047153745,0.00044349633,0.0007907893,0.0033834518,0.0012486962,0.0033264884,0.0014209545,0.001104646,0.00046503972],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007267113,0.0011080327,0.23928697,0.00075673335,0.00052063115,0.0005851386,0.006030521,0.19379483,0.05991667,0.04653995,0.003270259,0.4474636],"study_design_scores_gemma":[0.00008756667,0.002469682,0.112897426,0.00034200333,0.00027214846,0.00028292893,0.0038531707,0.8068075,0.033195574,0.0345635,0.0050110333,0.00021742206],"about_ca_topic_score_codex":0.0033122716,"about_ca_topic_score_gemma":0.0052286186,"teacher_disagreement_score":0.010902034,"about_ca_system_score_codex":0.002620702,"about_ca_system_score_gemma":0.0018704614,"threshold_uncertainty_score":0.05765617},"labels":[],"label_agreement":null},{"id":"W4226137778","doi":"10.1109/tse.2022.3201209","title":"Flakify: A Black-Box, Language Model-Based Predictor for Flaky Tests","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs; Canada Research Chairs; Western Canada Research Grid; Compute Canada","keywords":"Computer science; Debugging; Code coverage; Code (set theory); Set (abstract data type); Source code; White-box testing; Programming language; Software; Overhead (engineering); Black box; Test (biology); Machine learning; Artificial intelligence; Software development; Software construction","score_opus":0.013714965914711835,"score_gpt":0.24877966200787097,"score_spread":0.23506469609315914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226137778","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.089647606,0.0015246393,0.8467523,0.0011428255,0.00027746885,0.00042032055,0.003173421,0.054260246,0.0028011806],"genre_scores_gemma":[0.645858,0.0005903872,0.33384275,0.0009961568,0.00020989793,0.00087209354,0.008829081,0.0020843896,0.006717304],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99712914,0.0005796375,0.00018284681,0.0006614553,0.0010890242,0.00035796425],"domain_scores_gemma":[0.98890173,0.007022789,0.0015559274,0.0006268752,0.0013899634,0.00050266995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028134543,0.0030255343,0.0015329396,0.003036439,0.00049985107,0.0017852555,0.0019783408,0.0015846586,0.004328669],"category_scores_gemma":[0.017843857,0.00071288645,0.0013143352,0.00097393105,0.00093848124,0.002848697,0.0024710651,0.0038291665,0.0031327147],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018850644,0.001275659,0.07272731,0.00072494894,0.00033295137,0.0007492206,0.0003475074,0.3553134,0.025055079,0.0058262427,0.023767522,0.5119951],"study_design_scores_gemma":[0.00004852785,0.0002759219,0.0026703598,0.000057570243,0.000045198045,0.00009786372,0.000030690248,0.98430485,0.007010586,0.0030858025,0.0023254715,0.00004718559],"about_ca_topic_score_codex":0.007705452,"about_ca_topic_score_gemma":0.010190872,"teacher_disagreement_score":0.007705452,"about_ca_system_score_codex":0.0011848464,"about_ca_system_score_gemma":0.0035988155,"threshold_uncertainty_score":0.015321195},"labels":[],"label_agreement":null},{"id":"W4226205863","doi":"10.1109/qrs54544.2021.00076","title":"Vulnerability Analysis of Similar Code","year":2021,"lang":"en","type":"article","venue":"2021 IEEE 21st International Conference on Software Quality, Reliability and Security (QRS)","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"New York Institute of Technology","funders":"","keywords":"Computer science; Secure coding; Vulnerability (computing); Software security assurance; Code (set theory); Computer security; Code review; Static program analysis; Domain (mathematical analysis); Software; Scripting language; Source code; Vulnerability assessment; Security bug; Application security; Software engineering; Information security; Programming language; Software development; Security service; Mathematics","score_opus":0.060345565260924376,"score_gpt":0.3613439035872422,"score_spread":0.30099833832631784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226205863","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9770074,0.0008988419,0.014568059,0.00010048763,0.00002080208,0.00012524973,0.004723755,0.0007402314,0.0018151513],"genre_scores_gemma":[0.968457,0.00031201585,0.018559966,0.000055917943,0.000016922633,0.00012489021,0.011434686,0.00016390478,0.00087466126],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99619484,0.0004405602,0.00038801203,0.0011429752,0.0015223039,0.00031130714],"domain_scores_gemma":[0.9836083,0.006496385,0.0044019157,0.0023237506,0.002629205,0.0005404951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017562163,0.0005574684,0.0005310334,0.01084073,0.0007026822,0.0008326156,0.0005660453,0.0006371642,0.00076588034],"category_scores_gemma":[0.015771972,0.00022279735,0.00092790165,0.0051127565,0.0007145038,0.0016457448,0.0015346468,0.00055014057,0.00033560168],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039825888,0.00024639472,0.8178826,0.00087649166,0.00065287447,0.0016280371,0.0021684773,0.009976147,0.01966852,0.0037449547,0.005439931,0.13731723],"study_design_scores_gemma":[0.00003642501,0.0003821667,0.8486122,0.00021681043,0.00035560655,0.0052194814,0.0019406048,0.100133985,0.016026417,0.009109963,0.017839137,0.0001272298],"about_ca_topic_score_codex":0.004464858,"about_ca_topic_score_gemma":0.004423911,"teacher_disagreement_score":0.01084073,"about_ca_system_score_codex":0.0006273366,"about_ca_system_score_gemma":0.0007081648,"threshold_uncertainty_score":0.009287894},"labels":[],"label_agreement":null},{"id":"W4226253895","doi":"10.48550/arxiv.2202.03270","title":"Do Developers Refactor Data Access Code? An Empirical Study","year":2022,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Maintainability; Data access; Program comprehension; Code (set theory); Software; Software engineering; Database; Software system; Programming language","score_opus":0.35974523666241137,"score_gpt":0.3307371870697894,"score_spread":0.029008049592621987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226253895","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9967276,0.0004155297,0.0006865276,0.00060101325,0.000007341683,0.00009122111,0.00021816105,0.000020177802,0.0012323502],"genre_scores_gemma":[0.99725705,0.0004101989,0.001087403,0.00025396433,0.000016078386,0.0001346594,0.00027604503,0.000021786858,0.0005427789],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97619694,0.0095404815,0.0028225984,0.003007586,0.007020163,0.0014122914],"domain_scores_gemma":[0.44584474,0.40292957,0.08410711,0.016209505,0.045852765,0.0050562513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029929524,0.00042299795,0.00045554814,0.0042906236,0.0015135523,0.0023242647,0.0014969626,0.0017389884,0.0022520465],"category_scores_gemma":[0.23945211,0.0007925017,0.00035951825,0.0039597866,0.0022161554,0.005168572,0.0020331086,0.0024561698,0.0007639039],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022306711,0.0012135012,0.88893956,0.00045540943,0.00007264202,0.0006473598,0.06882352,0.00016011715,0.0006343477,0.00048902724,0.0019481649,0.03639324],"study_design_scores_gemma":[0.00009152448,0.0008219885,0.90996045,0.00056649605,0.00008208279,0.0014197989,0.07183409,0.0020722987,0.0010724975,0.0006609311,0.011343344,0.00007462591],"about_ca_topic_score_codex":0.006212248,"about_ca_topic_score_gemma":0.009050538,"teacher_disagreement_score":0.029929524,"about_ca_system_score_codex":0.0020872392,"about_ca_system_score_gemma":0.0027091082,"threshold_uncertainty_score":0.1582843},"labels":[],"label_agreement":null},{"id":"W4226265053","doi":"10.5281/zenodo.6338380","title":"Refactoring Debt: Myth or Reality? An Exploratory Study on the Relationship Between Technical Debt and Refactoring","year":2022,"lang":"en","type":"dataset","venue":"Espace ÉTS (ETS)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Code refactoring; Technical debt; Codebase; Computer science; Software engineering; Java; Source code; Code (set theory); Code smell; Debt; Software quality; Programming language; Software development; Software; Business; Finance","score_opus":0.11992321865469671,"score_gpt":0.35520456869998296,"score_spread":0.23528135004528625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226265053","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05828106,0.00051161245,0.0008618566,0.0010054297,0.00007038385,0.000104938874,0.93479604,0.0005355566,0.00383316],"genre_scores_gemma":[0.023168271,0.00014070324,0.0018534989,0.00016744783,0.000012706769,0.00023737478,0.9728252,0.00006547527,0.0015292278],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975922,0.0006361278,0.00036358123,0.00047130368,0.0007072077,0.00022966957],"domain_scores_gemma":[0.98793924,0.005150961,0.0016483971,0.002216691,0.0025288237,0.00051585416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002899745,0.00059421797,0.0004423298,0.0033111381,0.0009165286,0.0011167491,0.0014661351,0.001160534,0.0047247373],"category_scores_gemma":[0.013537483,0.00028569248,0.00055778597,0.0050816834,0.00034701978,0.0011330573,0.0018956142,0.0012377057,0.0050219526],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004525076,0.00042469,0.10542596,0.0016093755,0.000112150236,0.00037100664,0.0013634275,0.0011183788,0.0010386747,0.0019800689,0.850665,0.035438783],"study_design_scores_gemma":[0.0002814331,0.00014339584,0.24609646,0.00077887345,0.00009226775,0.000716113,0.0039838636,0.004356438,0.0022457521,0.002006061,0.73918825,0.00011111226],"about_ca_topic_score_codex":0.022403881,"about_ca_topic_score_gemma":0.063372724,"teacher_disagreement_score":0.022403881,"about_ca_system_score_codex":0.001448325,"about_ca_system_score_gemma":0.0013291156,"threshold_uncertainty_score":0.04454696},"labels":[],"label_agreement":null},{"id":"W4226265546","doi":"10.1145/3524842.3528435","title":"Does this apply to me?","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Canada Research Chairs","keywords":"Computer science; Context (archaeology); Focus (optics); Task (project management); Data science; World Wide Web; Information retrieval; Engineering","score_opus":0.01847715162260648,"score_gpt":0.2852877144144154,"score_spread":0.2668105627918089,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226265546","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09674425,0.01151789,0.009337508,0.543295,0.016516369,0.00015492906,0.0014157721,0.0008017159,0.32021657],"genre_scores_gemma":[0.52767134,0.00973662,0.005589751,0.15918687,0.0037479422,0.0001611491,0.00090841396,0.00077001745,0.2922279],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99806815,0.00078291417,0.00007256897,0.00027406105,0.00046462897,0.0003377132],"domain_scores_gemma":[0.994599,0.0017662172,0.0008250709,0.00024767453,0.0015282443,0.0010338449],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0024818953,0.0004259315,0.0003500024,0.0005414968,0.0031222978,0.003000863,0.0005001592,0.0023303076,0.04316609],"category_scores_gemma":[0.019015558,0.00015107391,0.0003836374,0.00055418164,0.002176318,0.004764423,0.0018955708,0.0029153577,0.018548045],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014291055,0.000094982475,0.021521494,0.0005794673,0.000046155958,0.0024996072,0.11136822,0.00006550232,0.0019458244,0.057805683,0.6795276,0.12440264],"study_design_scores_gemma":[0.000009253001,0.000056173074,0.006349232,0.00045440916,0.000015105865,0.0014932768,0.07076133,0.000120800665,0.00042703774,0.009033765,0.91122985,0.00004977385],"about_ca_topic_score_codex":0.0034617553,"about_ca_topic_score_gemma":0.0052941,"teacher_disagreement_score":0.9568339,"about_ca_system_score_codex":0.0011950362,"about_ca_system_score_gemma":0.0014190705,"threshold_uncertainty_score":0.144405},"labels":[],"label_agreement":null},{"id":"W4226304690","doi":"10.1145/3524610.3528387","title":"A first look at duplicate and near-duplicate self-admitted technical debt comments","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Debt; Information retrieval; Business; Finance","score_opus":0.019678405131037018,"score_gpt":0.27343940733900535,"score_spread":0.2537610022079683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226304690","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.565875,0.0149827,0.19908556,0.013309403,0.0036356663,0.00095742976,0.032645687,0.006928409,0.16258019],"genre_scores_gemma":[0.7580351,0.0043827477,0.13570587,0.0042060576,0.0013757647,0.00047024144,0.026055193,0.0056283996,0.06414048],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9816332,0.0031339575,0.0016146506,0.001922044,0.010432377,0.0012637414],"domain_scores_gemma":[0.783118,0.106808394,0.01926079,0.02423721,0.06453009,0.0020454722],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071637486,0.0006085621,0.0007152096,0.009930717,0.0034644147,0.004473788,0.0016503407,0.002074102,0.009446167],"category_scores_gemma":[0.10421671,0.0006644793,0.0006311595,0.012475269,0.001979419,0.008606934,0.004203014,0.003012412,0.00433454],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017261223,0.00047529017,0.13686806,0.0042384686,0.00028226827,0.011209236,0.06818298,0.0028650549,0.020788392,0.12141397,0.14767486,0.4842753],"study_design_scores_gemma":[0.00007422436,0.00023739203,0.09737994,0.002662099,0.0001430817,0.011819387,0.029917633,0.010961305,0.016039895,0.060399342,0.77006525,0.00030046995],"about_ca_topic_score_codex":0.00515989,"about_ca_topic_score_gemma":0.0061238026,"teacher_disagreement_score":0.009930717,"about_ca_system_score_codex":0.0020280695,"about_ca_system_score_gemma":0.00270764,"threshold_uncertainty_score":0.037885964},"labels":[],"label_agreement":null},{"id":"W4229000353","doi":"10.1145/3477314.3507053","title":"Fighting evil is not enough when refactoring metamodels","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 37th ACM/SIGAPP Symposium on Applied Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Code refactoring; Correctness; Computer science; Context (archaeology); Quality (philosophy); Task (project management); Process (computing); Set (abstract data type); Domain (mathematical analysis); Metamodeling; Software engineering; Heuristic; Artificial intelligence; Programming language; Systems engineering; Engineering; Software","score_opus":0.027770503131642404,"score_gpt":0.24767176647928307,"score_spread":0.21990126334764068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229000353","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5467642,0.0011049358,0.43133128,0.0020309046,0.00013197973,0.00039261565,0.00031213052,0.012852052,0.0050800345],"genre_scores_gemma":[0.627458,0.00027212102,0.36820573,0.0004563924,0.000029217314,0.000115402094,0.00055566017,0.0013325324,0.0015749321],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9909739,0.004329413,0.00051295164,0.00110677,0.0027168288,0.00036015152],"domain_scores_gemma":[0.9571003,0.029184401,0.0031037366,0.007524389,0.002380739,0.0007064846],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008394767,0.0015304498,0.0012079962,0.0010468639,0.00096471026,0.0028624758,0.0015576696,0.0021852176,0.001645466],"category_scores_gemma":[0.03904507,0.0008801713,0.0011219494,0.00070072315,0.0009059536,0.003965625,0.001975512,0.0020385794,0.0008615646],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030535366,0.0018514296,0.02784678,0.0017227163,0.0006466602,0.0012300097,0.004263402,0.09477236,0.20272507,0.0074713975,0.008791816,0.6456249],"study_design_scores_gemma":[0.00053167366,0.0030249958,0.022406483,0.00053163257,0.00060414436,0.0020682812,0.0030730865,0.74733406,0.15616627,0.03465644,0.029273758,0.00032926435],"about_ca_topic_score_codex":0.0014174923,"about_ca_topic_score_gemma":0.0030955062,"teacher_disagreement_score":0.008394767,"about_ca_system_score_codex":0.00067500956,"about_ca_system_score_gemma":0.0012865303,"threshold_uncertainty_score":0.04439628},"labels":[],"label_agreement":null},{"id":"W4229442351","doi":"10.26226/morressier.613b5418842293c031b5b5dd","title":"Understanding Quantum Software Engineering Challenges: An Empirical Study on Stack Exchange Forums and GitHub Issues","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Stack (abstract data type); Software; Computer science; Empirical research; Software engineering; Quantum; Data science; Operating system; Physics; Epistemology","score_opus":0.2574576734917611,"score_gpt":0.37103605524958805,"score_spread":0.11357838175782697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229442351","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.993252,0.00014222751,0.0008379779,0.00063286175,0.000018684534,0.00007295832,0.0001490682,0.000061391256,0.0048328172],"genre_scores_gemma":[0.9964302,0.00015340447,0.00089317927,0.00020533308,0.000044345026,0.00010483118,0.00041638679,0.00009986694,0.0016524738],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99143773,0.004300374,0.0004939336,0.000657934,0.0022225305,0.00088750606],"domain_scores_gemma":[0.8383723,0.10880325,0.024041273,0.0054654893,0.014392849,0.008924863],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013666108,0.00040367967,0.0005957394,0.004683267,0.0044592936,0.006341971,0.0011477367,0.0022186253,0.0048975865],"category_scores_gemma":[0.10893592,0.00043951411,0.00031063714,0.004875913,0.0023339775,0.012841218,0.004361299,0.0028307766,0.0015375796],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089177425,0.0034528924,0.6037087,0.00050958904,0.000115769035,0.0013912064,0.26712215,0.00071626727,0.002383365,0.012113857,0.013114839,0.09447959],"study_design_scores_gemma":[0.0001142425,0.00080683426,0.5843563,0.00044933692,0.00012632366,0.000804355,0.34748453,0.011089258,0.0015957837,0.009850466,0.04311897,0.00020354072],"about_ca_topic_score_codex":0.008079914,"about_ca_topic_score_gemma":0.0069674966,"teacher_disagreement_score":0.9863339,"about_ca_system_score_codex":0.0020134777,"about_ca_system_score_gemma":0.002000664,"threshold_uncertainty_score":0.07227415},"labels":[],"label_agreement":null},{"id":"W4229517920","doi":"10.26226/morressier.613b54401459512fce6a7cfd","title":"Leveraging Unsupervised Learning to Summarize APIs Discussed in Stack Overflow","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; École de Technologie Supérieure","funders":"","keywords":"Stack (abstract data type); Computer science; Unsupervised learning; Artificial intelligence; Machine learning; Operating system","score_opus":0.028388230463872337,"score_gpt":0.2734209649122384,"score_spread":0.24503273444836604,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229517920","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26512453,0.0046089115,0.65452087,0.0020254964,0.0010024257,0.00043833436,0.012613682,0.050226297,0.009439433],"genre_scores_gemma":[0.665759,0.00130714,0.28489023,0.0004774717,0.0008010962,0.00035887692,0.03473212,0.0030825082,0.008591438],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99862945,0.00023104605,0.00011976513,0.0004684567,0.000381454,0.00016991599],"domain_scores_gemma":[0.99514186,0.0019537078,0.0006661778,0.0008094485,0.0011725505,0.0002562707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013727334,0.0013311197,0.0008604624,0.006396304,0.00093759864,0.0022194204,0.0013226407,0.0011568349,0.0018227111],"category_scores_gemma":[0.009514807,0.0004672904,0.0011720692,0.003095578,0.00054118276,0.004151013,0.0020983436,0.0017437224,0.0022619285],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091938523,0.00062628323,0.039095215,0.0006124076,0.0003683471,0.0005789351,0.0009767914,0.02743938,0.023812763,0.004390084,0.060865127,0.8403154],"study_design_scores_gemma":[0.000062137326,0.00022923501,0.016697537,0.000106829204,0.00027718768,0.00027699454,0.00042372907,0.9096167,0.01571059,0.031018768,0.025464185,0.00011616013],"about_ca_topic_score_codex":0.0052443114,"about_ca_topic_score_gemma":0.012124834,"teacher_disagreement_score":0.006396304,"about_ca_system_score_codex":0.0007598983,"about_ca_system_score_gemma":0.0014816519,"threshold_uncertainty_score":0.010427594},"labels":[],"label_agreement":null},{"id":"W4229596932","doi":"10.1109/icse.2003.1201189","title":"Using benchmarking to advance research: a challenge to software engineering","year":2003,"lang":"en","type":"article","venue":"25th International Conference on Software Engineering, 2003. Proceedings.","topic":"Software Engineering Research","field":"Computer Science","cited_by":150,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Benchmarking; Benchmark (surveying); Computer science; Data science; Software engineering; Software; Academic community; Engineering management; Management science; Engineering","score_opus":0.11012686057636706,"score_gpt":0.3553565504832152,"score_spread":0.24522968990684813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229596932","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018883402,0.028022975,0.6365957,0.28054297,0.0042177844,0.00062923215,0.00033181533,0.0037208863,0.027055277],"genre_scores_gemma":[0.20736398,0.015213503,0.7508223,0.014939757,0.0035339652,0.0020416898,0.00077916763,0.0024742517,0.0028313866],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.6199469,0.25066832,0.02084692,0.015539281,0.088458814,0.004539864],"domain_scores_gemma":[0.2921697,0.4562198,0.027531328,0.11629769,0.09335408,0.014427461],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.30673128,0.0033124432,0.0054695522,0.012507649,0.006122894,0.030853327,0.01107002,0.009709315,0.0019050738],"category_scores_gemma":[0.52741086,0.0014535624,0.0017279002,0.018374143,0.021529878,0.05805392,0.02213132,0.018721705,0.0024956746],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012503231,0.00035741655,0.0055687805,0.0019192201,0.00018523783,0.00017150791,0.0044922694,0.009868663,0.0013822116,0.45449317,0.03955973,0.4818767],"study_design_scores_gemma":[0.00008790466,0.0003992956,0.0015056969,0.002298149,0.000049074817,0.00016942837,0.0033330452,0.014983171,0.00163817,0.86365235,0.11167177,0.00021193683],"about_ca_topic_score_codex":0.0027564478,"about_ca_topic_score_gemma":0.0025146632,"teacher_disagreement_score":0.6932687,"about_ca_system_score_codex":0.007025863,"about_ca_system_score_gemma":0.024001725,"threshold_uncertainty_score":0.8549237},"labels":[],"label_agreement":null},{"id":"W4229937035","doi":"10.32920/ryerson.14668263","title":"Researchgate.net crawler and a new contribution determines sequence (CDS) method","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Web crawler; Crawling; Scripting language; Focused crawler; Computer science; World Wide Web; Field (mathematics); Java; Information retrieval; Data mining; Data science; The Internet; Web server; Operating system; Static web page; Mathematics","score_opus":0.09491528026853427,"score_gpt":0.3891979685653524,"score_spread":0.29428268829681814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229937035","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022547405,0.0011166211,0.85854876,0.000742081,0.0005572446,0.001701369,0.017507792,0.08375985,0.013518895],"genre_scores_gemma":[0.04776071,0.00040012304,0.91548795,0.00012169325,0.00013605584,0.0010617373,0.017169088,0.005586522,0.012276128],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9952154,0.00093351625,0.0007096022,0.0012035302,0.0017681387,0.00016974543],"domain_scores_gemma":[0.9920763,0.0032285678,0.00046814742,0.0015209409,0.00237523,0.00033086457],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038544703,0.0015301746,0.001320596,0.014347897,0.0010181516,0.0030215497,0.0014272382,0.0014041221,0.008224028],"category_scores_gemma":[0.02053499,0.0010242581,0.0013362896,0.008686018,0.0006155015,0.0032483123,0.001927129,0.0011354802,0.0055214465],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005574642,0.00031134882,0.013140412,0.00177803,0.00040752022,0.00043869202,0.0008898772,0.009458526,0.009064628,0.031068306,0.10530045,0.8275847],"study_design_scores_gemma":[0.00072005077,0.00034224562,0.014109539,0.00032689652,0.00039741912,0.0016809563,0.000490335,0.41936758,0.03220912,0.03923316,0.49084386,0.0002788002],"about_ca_topic_score_codex":0.008970808,"about_ca_topic_score_gemma":0.013183912,"teacher_disagreement_score":0.014347897,"about_ca_system_score_codex":0.001391957,"about_ca_system_score_gemma":0.0040285443,"threshold_uncertainty_score":0.027512133},"labels":[],"label_agreement":null},{"id":"W4230282342","doi":"10.32920/ryerson.14660433","title":"Predicting the time-to-deliver of software changes","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software; Data mining; Identification (biology); Event (particle physics); Feature (linguistics); Data science","score_opus":0.02067907124096478,"score_gpt":0.25595874584334344,"score_spread":0.23527967460237867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230282342","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91503507,0.00065464573,0.06849921,0.00087050174,0.00010864753,0.00014888844,0.011654763,0.0007768226,0.002251385],"genre_scores_gemma":[0.96289086,0.00042947047,0.022410752,0.000082634186,0.00007586889,0.00013629923,0.0117967455,0.00007808189,0.0020993652],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9983041,0.00041763618,0.00016281834,0.0004594121,0.00050735974,0.0001486833],"domain_scores_gemma":[0.9397877,0.041068714,0.010319922,0.0037511927,0.0038037102,0.0012688002],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004773933,0.0006404128,0.0005134913,0.0031235341,0.0002640827,0.0012103948,0.00079939957,0.0013710303,0.0025711958],"category_scores_gemma":[0.055679746,0.00029651125,0.0006750394,0.0024200964,0.00030866897,0.001806582,0.0006568018,0.0016618052,0.0018155875],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004680788,0.0004776906,0.74086624,0.0003936827,0.000192316,0.00025571254,0.00056604174,0.108918995,0.0030330364,0.0053631915,0.004911493,0.13455355],"study_design_scores_gemma":[0.000044832454,0.0009049158,0.48454747,0.00014066728,0.00013157804,0.000389272,0.0006961826,0.485274,0.005954447,0.012241619,0.009570254,0.00010483655],"about_ca_topic_score_codex":0.0042277114,"about_ca_topic_score_gemma":0.0066036675,"teacher_disagreement_score":0.004773933,"about_ca_system_score_codex":0.0006780945,"about_ca_system_score_gemma":0.00063022634,"threshold_uncertainty_score":0.025247276},"labels":[],"label_agreement":null},{"id":"W4230524485","doi":"10.32920/ryerson.14646261","title":"On Predicting Rediscoveries of Software Defects","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Eclipse; Downtime; Computer science; Software bug; Forcing (mathematics); Customer satisfaction; Software; Quality (philosophy); Software quality; Software engineering; Data science; Software development; Operating system; Business; Marketing","score_opus":0.018148737736127903,"score_gpt":0.26380164501338726,"score_spread":0.24565290727725936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230524485","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95857227,0.0026182171,0.025302552,0.0016206756,0.0000904725,0.000118086195,0.008348831,0.0013014411,0.0020273898],"genre_scores_gemma":[0.96301633,0.0008078575,0.018566966,0.00017775816,0.00009101616,0.00006304697,0.015874976,0.000048754297,0.001353292],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984861,0.00033460037,0.00013595527,0.0005418895,0.00034198485,0.00015952412],"domain_scores_gemma":[0.9804176,0.012965155,0.0024688118,0.0012927646,0.0022425002,0.0006130719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031539071,0.0011415357,0.0007467812,0.0053547486,0.0004795182,0.0012319373,0.0013241142,0.0019212398,0.00082984107],"category_scores_gemma":[0.01764614,0.0004871106,0.00089956645,0.0037812553,0.0003548192,0.001949881,0.00073405117,0.0017675346,0.0008166902],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033592858,0.00080561946,0.7796904,0.00027150844,0.00031489463,0.00040289218,0.0003394501,0.086298294,0.0012454216,0.0008752359,0.011109525,0.11831078],"study_design_scores_gemma":[0.00004260467,0.00026440967,0.18706237,0.00009020794,0.00017014684,0.00040722208,0.00036177223,0.80493426,0.0012940695,0.001813078,0.0035071443,0.000052784235],"about_ca_topic_score_codex":0.049042627,"about_ca_topic_score_gemma":0.06764831,"teacher_disagreement_score":0.049042627,"about_ca_system_score_codex":0.0008776104,"about_ca_system_score_gemma":0.0009606993,"threshold_uncertainty_score":0.09751439},"labels":[],"label_agreement":null},{"id":"W4230584787","doi":"10.1007/978-0-387-30164-8_661","title":"Predictive Techniques in Software Engineering","year":2010,"lang":"en","type":"book-chapter","venue":"Encyclopedia of Machine Learning","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Software engineering","score_opus":0.005930212369387834,"score_gpt":0.22259684999099674,"score_spread":0.2166666376216089,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230584787","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025653085,0.119325176,0.7081272,0.0029017776,0.0016896622,0.00007040947,0.00027826062,0.0020598597,0.16298233],"genre_scores_gemma":[0.13139293,0.16777217,0.38439772,0.0015554221,0.004618613,0.00030706913,0.0017115519,0.0010607953,0.3071837],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99949706,0.00007798371,0.000021562413,0.00008465599,0.0003016021,0.000017138591],"domain_scores_gemma":[0.9991302,0.0005639948,0.000030445251,0.00011920582,0.0001355303,0.000020646627],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005147834,0.0010045948,0.00081528554,0.0011029836,0.0003868653,0.0018292249,0.0014289205,0.00081420364,0.013245483],"category_scores_gemma":[0.0022218167,0.00045101173,0.00045520742,0.0022216528,0.0010663129,0.0027202892,0.0008860341,0.002279708,0.006817659],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019667717,0.00006203249,0.00022109058,0.0005816379,0.000032895543,0.00006826477,0.00015367164,0.014180647,0.0010693398,0.2883503,0.0729576,0.62230295],"study_design_scores_gemma":[0.000011521685,0.000050067203,0.00053908484,0.0005329456,0.000036799276,0.00038731616,0.000045739518,0.05996273,0.0020085336,0.5862881,0.3501034,0.000033711607],"about_ca_topic_score_codex":0.00079247385,"about_ca_topic_score_gemma":0.0010063066,"teacher_disagreement_score":0.013245483,"about_ca_system_score_codex":0.0005199765,"about_ca_system_score_gemma":0.0007232867,"threshold_uncertainty_score":0.04431057},"labels":[],"label_agreement":null},{"id":"W4230681096","doi":"10.1109/icse.2013.6606709","title":"V:Issue:lizer: Exploring requirements clarification in online communication over time","year":2013,"lang":"en","type":"article","venue":"2013 35th International Conference on Software Engineering (ICSE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Domain (mathematical analysis); World Wide Web; Software engineering","score_opus":0.08203369238501862,"score_gpt":0.3074966664419672,"score_spread":0.2254629740569486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230681096","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20036592,0.001215888,0.4071762,0.008409037,0.0013737078,0.0027532529,0.026234638,0.15238707,0.20008422],"genre_scores_gemma":[0.37770697,0.000938448,0.4459689,0.0018072195,0.00036283798,0.0021502783,0.021886168,0.018019088,0.1311601],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.9993793,0.00026511675,0.00002673532,0.000076460514,0.00019216856,0.000060207734],"domain_scores_gemma":[0.99481344,0.0038495245,0.00018089522,0.00038788823,0.0004775618,0.00029074075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016396021,0.00065157306,0.0004056811,0.0015378682,0.00074069866,0.0017911138,0.0009544625,0.0013706028,0.06425824],"category_scores_gemma":[0.008293119,0.0002785101,0.0005290789,0.0007679306,0.00035620405,0.004152983,0.0023599188,0.000876874,0.010604305],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008068922,0.00051130564,0.0056035435,0.0015485534,0.00006578853,0.0015091718,0.024248041,0.001839585,0.027787458,0.009014625,0.5068158,0.4202492],"study_design_scores_gemma":[0.00025488425,0.000826313,0.025870854,0.0009880287,0.00008169213,0.0019256165,0.016256295,0.044883806,0.024917535,0.01500266,0.86865526,0.00033707122],"about_ca_topic_score_codex":0.0014294193,"about_ca_topic_score_gemma":0.0036864355,"teacher_disagreement_score":0.06425824,"about_ca_system_score_codex":0.0004602212,"about_ca_system_score_gemma":0.0004117309,"threshold_uncertainty_score":0.21496522},"labels":[],"label_agreement":null},{"id":"W4231026759","doi":"10.1109/msr.2015.17","title":"An Empirical Study of the Copy and Paste Behavior during Development","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Copying; Computer science; Eclipse; Context (archaeology); Code (set theory); Programming language; Source code; Cloning (programming)","score_opus":0.05693587498825337,"score_gpt":0.34075599426623787,"score_spread":0.2838201192779845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231026759","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9968376,0.00023412336,0.0012474029,0.00010643062,0.00000456781,0.000033665874,0.0002914043,0.000034234865,0.0012106441],"genre_scores_gemma":[0.99595946,0.00034168825,0.0023266652,0.00006514963,0.000009835728,0.00006169501,0.0005806748,0.000032088465,0.00062272005],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.990585,0.0036399935,0.0010050219,0.0013936318,0.0029767838,0.00039952563],"domain_scores_gemma":[0.7949689,0.14027974,0.036121406,0.009417213,0.016707705,0.002505084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064696386,0.00031967077,0.0002759234,0.003136672,0.0006874953,0.0013449044,0.0007001714,0.0006987486,0.0009041658],"category_scores_gemma":[0.074407116,0.00048056082,0.0002054088,0.0031544454,0.0008447579,0.0027470004,0.0009633559,0.00081640505,0.00046359305],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001508631,0.00028862018,0.92436767,0.00034764505,0.000081624086,0.00035925853,0.019450454,0.0002814129,0.0029116946,0.0003342086,0.0013428116,0.050083764],"study_design_scores_gemma":[0.000013947464,0.00037776385,0.9732635,0.00015604614,0.000046384237,0.0011042932,0.012657385,0.0022029718,0.0026976566,0.00033133157,0.0071023474,0.00004633154],"about_ca_topic_score_codex":0.0019669086,"about_ca_topic_score_gemma":0.004065472,"teacher_disagreement_score":0.0064696386,"about_ca_system_score_codex":0.00063223933,"about_ca_system_score_gemma":0.0006470572,"threshold_uncertainty_score":0.034215093},"labels":[],"label_agreement":null},{"id":"W4231171191","doi":"10.1002/0471028959.sof110","title":"Experience Factory","year":2002,"lang":"en","type":"other","venue":"Encyclopedia of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"PricewaterhouseCoopers (Canada)","funders":"","keywords":"Factory (object-oriented programming); Reuse; Quality (philosophy); Process management; Product (mathematics); Product lifecycle; Computer science; Engineering management; Knowledge management; New product development; Engineering; Manufacturing engineering; Operations management; Business; Marketing; Waste management","score_opus":0.011001754170537503,"score_gpt":0.22699145970655,"score_spread":0.2159897055360125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231171191","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019870976,0.0007120107,0.308287,0.0035592304,0.00060727017,0.0009264283,0.002399338,0.01971227,0.6439255],"genre_scores_gemma":[0.28121787,0.001807821,0.23232871,0.0024934479,0.00039076834,0.0014969785,0.00988441,0.005450207,0.46492982],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99645835,0.0010153041,0.00025831172,0.000652109,0.0010715376,0.0005443236],"domain_scores_gemma":[0.99278533,0.0010657649,0.00039616856,0.0032499114,0.0014657036,0.0010371399],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027616601,0.00088237727,0.00051705557,0.0015944768,0.0013700534,0.007933133,0.002672706,0.0013249435,0.07696751],"category_scores_gemma":[0.0083669275,0.00061545457,0.00090531894,0.001099187,0.0015061169,0.007324897,0.0073475265,0.0019166433,0.026172398],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045155347,0.0005205772,0.004166075,0.00042038827,0.000040343126,0.00093943154,0.0065811016,0.0021768468,0.0038117324,0.29071105,0.15766722,0.5325137],"study_design_scores_gemma":[0.000047201156,0.00013530561,0.0013554398,0.00015834752,0.000019838824,0.00066182355,0.0011592302,0.002359296,0.0026293632,0.020763403,0.9706613,0.000049457656],"about_ca_topic_score_codex":0.0032126354,"about_ca_topic_score_gemma":0.0017797542,"teacher_disagreement_score":0.07696751,"about_ca_system_score_codex":0.0018781519,"about_ca_system_score_gemma":0.003935424,"threshold_uncertainty_score":0.257482},"labels":[],"label_agreement":null},{"id":"W4231242876","doi":"10.1145/3089649.3089655","title":"4th International Workshop on Conducting Empirical Studies in Industry (CESI 2016)","year":2017,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Empirical research; Software; Field (mathematics); Computer science; Replication (statistics); Software engineering; Engineering management; Engineering; Knowledge management; Management science; Mathematics","score_opus":0.196112487251442,"score_gpt":0.4040139648779832,"score_spread":0.20790147762654118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231242876","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007143599,0.07808853,0.4674288,0.17014758,0.06139359,0.0045232605,0.010449129,0.01249895,0.18832658],"genre_scores_gemma":[0.060655516,0.063951336,0.58938843,0.03630474,0.02229048,0.010979397,0.034248427,0.012351638,0.16983004],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.92163163,0.04969614,0.006007898,0.005131988,0.014432953,0.0030995277],"domain_scores_gemma":[0.8034068,0.08163796,0.005581209,0.04333459,0.04547579,0.020563688],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.14506955,0.002808765,0.0032008034,0.008297965,0.0030664473,0.016407827,0.00597944,0.007750867,0.07593121],"category_scores_gemma":[0.14597759,0.002197046,0.003517592,0.00591379,0.0054489695,0.013970332,0.022407854,0.011485855,0.04260084],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035150192,0.00046681377,0.0019429885,0.001665609,0.00014492984,0.0002533986,0.0030471582,0.0011011648,0.0025271203,0.05006918,0.5830343,0.35539594],"study_design_scores_gemma":[0.00005986058,0.000108132605,0.0021780992,0.00327806,0.0000352281,0.0002056369,0.0013652665,0.0007570256,0.00089981314,0.03474324,0.95630693,0.000062641106],"about_ca_topic_score_codex":0.0030606745,"about_ca_topic_score_gemma":0.0039713,"teacher_disagreement_score":0.85493046,"about_ca_system_score_codex":0.0043160343,"about_ca_system_score_gemma":0.018195296,"threshold_uncertainty_score":0.7672103},"labels":[],"label_agreement":null},{"id":"W4231296399","doi":"10.1007/978-3-540-78195-0_16","title":"Quantitative Approaches in Object-Oriented Software Engineering","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Session (web analytics); Computer science; Object (grammar); Closing (real estate); Software; Quality (philosophy); Metric (unit); Data science; Software engineering; Engineering; World Wide Web; Artificial intelligence; Epistemology; Political science; Operations management","score_opus":0.039633547060391586,"score_gpt":0.2562359961047838,"score_spread":0.21660244904439224,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231296399","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018846514,0.010163854,0.95835096,0.0014817861,0.00038516376,0.00006631803,0.00009764271,0.0003397885,0.027229758],"genre_scores_gemma":[0.12755047,0.018438589,0.8121339,0.00091491267,0.0013003661,0.0006024681,0.00034936852,0.00054906745,0.038160898],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99668556,0.0010878713,0.0002232249,0.00026006254,0.0016564543,0.00008672559],"domain_scores_gemma":[0.99322444,0.005306244,0.00024193642,0.00050225575,0.00063095556,0.000094179544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038264026,0.0012070438,0.0010786072,0.0027008513,0.00081701984,0.0029973027,0.0020125762,0.00085751334,0.006276504],"category_scores_gemma":[0.009999759,0.0009425387,0.0007053549,0.003451335,0.003912305,0.0057475124,0.0014140806,0.0026408364,0.0012786442],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000005604083,0.000030400797,0.00010409647,0.0003306567,0.000011941372,0.000019298617,0.00024361689,0.0027173373,0.00041832542,0.94251734,0.0037889504,0.049812283],"study_design_scores_gemma":[0.000005042668,0.000010935314,0.00013220636,0.000100348705,0.00000971236,0.000029958974,0.00006470367,0.0065713394,0.00031088846,0.9700787,0.02267665,0.00000947392],"about_ca_topic_score_codex":0.0010786785,"about_ca_topic_score_gemma":0.0012015977,"teacher_disagreement_score":0.006276504,"about_ca_system_score_codex":0.0017855979,"about_ca_system_score_gemma":0.0012165033,"threshold_uncertainty_score":0.020997047},"labels":[],"label_agreement":null},{"id":"W4231546496","doi":"10.1109/ms.2011.59","title":"Point/Counterpoint","year":2011,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Software engineering; Computer science; Software construction; Software development; Social software engineering; Predictability; Software; Architectural pattern; Component-based software engineering; Resource-oriented architecture; Software design; Software peer review; Programming language","score_opus":0.031384596718333616,"score_gpt":0.24662010987626712,"score_spread":0.2152355131579335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231546496","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013750922,0.0029328025,0.19561422,0.023553485,0.014223318,0.0011257263,0.0010799688,0.004062511,0.7436572],"genre_scores_gemma":[0.43016145,0.0029076117,0.11478316,0.014260117,0.0026834116,0.0013353741,0.0014023405,0.0019858722,0.43048075],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99610597,0.00076629757,0.00023096695,0.0011040475,0.0013206953,0.00047199964],"domain_scores_gemma":[0.9957403,0.0012336359,0.0002733306,0.0011530641,0.0012625968,0.00033705521],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025052452,0.0015378022,0.0007192869,0.0020627514,0.0027059026,0.0051972466,0.0025346423,0.0043243,0.15234451],"category_scores_gemma":[0.0156788,0.0005055356,0.0015041521,0.0010565615,0.005047042,0.008292104,0.0049629137,0.0042413855,0.04020007],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003535613,0.000067736735,0.0010739763,0.00030157657,0.00003733486,0.001621201,0.0010807446,0.00069988583,0.0016750309,0.8692198,0.06477125,0.059097867],"study_design_scores_gemma":[0.00009217569,0.00014507564,0.0003494908,0.00035867174,0.00004854312,0.0014490633,0.0009641076,0.0018165944,0.0028772098,0.2362812,0.755564,0.000053953543],"about_ca_topic_score_codex":0.002651763,"about_ca_topic_score_gemma":0.0015578342,"teacher_disagreement_score":0.15234451,"about_ca_system_score_codex":0.001971101,"about_ca_system_score_gemma":0.0020517441,"threshold_uncertainty_score":0.5096432},"labels":[],"label_agreement":null},{"id":"W4231615358","doi":"10.1109/.2001.914966","title":"Cohesion as changeability indicator in object-oriented systems","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Cohesion (chemistry); Computer science; Object-oriented programming; Software; Realm; Reliability engineering; Industrial engineering; Engineering; Programming language","score_opus":0.023386611698478145,"score_gpt":0.2550835652323642,"score_spread":0.23169695353388606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231615358","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86421156,0.0018901936,0.12538865,0.00034503572,0.00006846209,0.0001673237,0.00021479807,0.00073444657,0.006979415],"genre_scores_gemma":[0.9898517,0.00016981563,0.009419184,0.00001924207,0.000034548073,0.00004937912,0.00015887538,0.000036511377,0.00026083665],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993111,0.0024002166,0.00055132876,0.00069967157,0.002942924,0.0002948434],"domain_scores_gemma":[0.94084734,0.039545186,0.011041243,0.0032554572,0.0043893675,0.00092138944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044706017,0.0004367955,0.0006044941,0.0073525608,0.0006273145,0.0017265527,0.0004603693,0.0007555896,0.00061366544],"category_scores_gemma":[0.04996252,0.0003473536,0.00040340424,0.0038897975,0.0018457316,0.0030144618,0.0012625373,0.00067742344,0.00013664842],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078132143,0.00056573347,0.45892668,0.0012257139,0.00041054416,0.0011983786,0.011709907,0.07276013,0.04453305,0.058448013,0.0027275737,0.34671292],"study_design_scores_gemma":[0.00009644573,0.0020419382,0.6623627,0.0002846027,0.0003232552,0.0015905664,0.0027227453,0.21874405,0.024488674,0.07792318,0.009162039,0.00025989834],"about_ca_topic_score_codex":0.0014420231,"about_ca_topic_score_gemma":0.0007166113,"teacher_disagreement_score":0.0073525608,"about_ca_system_score_codex":0.0009839222,"about_ca_system_score_gemma":0.0003949525,"threshold_uncertainty_score":0.023643076},"labels":[],"label_agreement":null},{"id":"W4231932677","doi":"10.1145/1095430.1081761","title":"Anchoring and adjustment in software estimation","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Anchoring; Estimator; Respondent; Estimation; Software; Computer science; Econometrics; Statistics; Psychology; Mathematics; Social psychology; Engineering","score_opus":0.015244430479361889,"score_gpt":0.2533436695576924,"score_spread":0.2380992390783305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231932677","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36209115,0.006304077,0.5611194,0.007415217,0.0007108039,0.0007605085,0.0001653494,0.0006451029,0.060788427],"genre_scores_gemma":[0.93684053,0.0010661845,0.05934284,0.00080266927,0.00027659416,0.00030718106,0.00006152263,0.00016207756,0.0011403303],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.80529994,0.14273842,0.009864068,0.012041127,0.027578684,0.0024777425],"domain_scores_gemma":[0.29130468,0.6037305,0.049530104,0.036619622,0.017369084,0.0014460527],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.100228064,0.0011764617,0.0010831392,0.0034851031,0.0022898798,0.00742431,0.0017804514,0.0037504057,0.0046705436],"category_scores_gemma":[0.6326747,0.0015458158,0.0014150005,0.0046864687,0.01104954,0.012001254,0.0061770543,0.005699454,0.0008737185],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001991735,0.00043500587,0.1309579,0.0017250415,0.0012508566,0.00066999864,0.056933206,0.026231756,0.007359875,0.28953022,0.0041939905,0.4787204],"study_design_scores_gemma":[0.00041335,0.001163325,0.13971367,0.0016571633,0.00082105125,0.001063816,0.010902159,0.059995964,0.012551748,0.7359032,0.034901585,0.0009128748],"about_ca_topic_score_codex":0.0039643277,"about_ca_topic_score_gemma":0.002107991,"teacher_disagreement_score":0.100228064,"about_ca_system_score_codex":0.0029421076,"about_ca_system_score_gemma":0.0020365461,"threshold_uncertainty_score":0.53006303},"labels":[],"label_agreement":null},{"id":"W4232217129","doi":"10.1007/978-3-030-24094-3_2","title":"Encapsulation","year":2019,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Encapsulation (networking); Software; Software engineering; Programming language","score_opus":0.021941026214444187,"score_gpt":0.2427677840133828,"score_spread":0.22082675779893862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232217129","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020460233,0.0006256223,0.12748538,0.000646937,0.0008161482,0.000207795,0.0013291342,0.0073914877,0.8594515],"genre_scores_gemma":[0.028855659,0.0012593108,0.03388822,0.00066744915,0.0002634227,0.00019895134,0.002927117,0.00441288,0.9275271],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99922764,0.00009336323,0.000037434453,0.00015854019,0.00037323203,0.00010981758],"domain_scores_gemma":[0.9992435,0.000106061765,0.000025124147,0.00041455496,0.00017825606,0.000032611224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005981497,0.0012987495,0.00055058853,0.0012919983,0.0010762269,0.0042687776,0.0015976522,0.0009776526,0.16193059],"category_scores_gemma":[0.0018992743,0.00066083844,0.0008658878,0.0011272561,0.0011001796,0.0051821875,0.0031603512,0.001977939,0.12547609],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000063216554,0.000055740016,0.000108583954,0.0001953084,0.00000961866,0.000095553885,0.000430881,0.0004453668,0.006653718,0.5463982,0.19709916,0.24844465],"study_design_scores_gemma":[0.0000076498345,0.000014369233,0.000074559575,0.00006796292,0.00001128773,0.0002063753,0.00006470377,0.00074572983,0.006415716,0.05178796,0.9405892,0.00001446741],"about_ca_topic_score_codex":0.0013882264,"about_ca_topic_score_gemma":0.0012921158,"teacher_disagreement_score":0.16193059,"about_ca_system_score_codex":0.0008740953,"about_ca_system_score_gemma":0.00092658604,"threshold_uncertainty_score":0.5417118},"labels":[],"label_agreement":null},{"id":"W4232383598","doi":"10.4018/9781591409411.ch002.ch000","title":"Intelligent Analysis of Software Maintenance Data","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software; Software maintenance; Software development; Software engineering; Process (computing); Data mining; Data extraction; Software development process; Software construction; Software sizing; Reliability engineering; Engineering","score_opus":0.055673288942141685,"score_gpt":0.284304236008936,"score_spread":0.22863094706679432,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232383598","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26559025,0.0014448696,0.7085358,0.00083257153,0.000087438755,0.00047878936,0.006726465,0.0103806155,0.005923248],"genre_scores_gemma":[0.52356005,0.0007390547,0.4601282,0.00012296873,0.00006031054,0.000283985,0.013002163,0.00022880417,0.0018745281],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99766195,0.00042998855,0.00025564214,0.00055203965,0.0009986309,0.00010169897],"domain_scores_gemma":[0.9943229,0.0027601658,0.0005934994,0.00081153674,0.0014262078,0.000085701016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002511732,0.00057061994,0.0009811013,0.0058769584,0.0004475349,0.0017594678,0.0007760149,0.00052879355,0.0009137218],"category_scores_gemma":[0.010343893,0.00031552804,0.0009010651,0.0029079376,0.00022454286,0.0014532738,0.0009896802,0.00081719976,0.0005676789],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030488102,0.00018916489,0.035897925,0.0004143638,0.0002011312,0.00037434115,0.000706955,0.059671994,0.017024197,0.004327313,0.005083139,0.8758045],"study_design_scores_gemma":[0.000017855935,0.00014437365,0.03761919,0.000109523025,0.000108654865,0.0002909952,0.0003431884,0.91818875,0.019561041,0.010416305,0.0131488005,0.00005139722],"about_ca_topic_score_codex":0.00229749,"about_ca_topic_score_gemma":0.0032119497,"teacher_disagreement_score":0.0058769584,"about_ca_system_score_codex":0.0007111863,"about_ca_system_score_gemma":0.0008393339,"threshold_uncertainty_score":0.0132834315},"labels":[],"label_agreement":null},{"id":"W4232396147","doi":"10.1109/issre.2007.36","title":"Using Machine Learning to Support Debugging with Tarantula","year":2007,"lang":"en","type":"article","venue":"Proceedings/Proceedings - International Symposium on Software Reliability Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Debugging; Computer science; Statement (logic); Ranking (information retrieval); Test (biology); Software bug; Decision tree; Fault (geology); Machine learning; Fault tree analysis; Test case; Algorithmic program debugging; Artificial intelligence; Tree (set theory); Reliability engineering; Programming language; Software; Engineering","score_opus":0.014382982750755039,"score_gpt":0.2633939116122934,"score_spread":0.2490109288615384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232396147","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021358054,0.000082041865,0.96449983,0.00020115169,0.000023128703,0.000084189196,0.00010781773,0.012876084,0.0007676535],"genre_scores_gemma":[0.19623132,0.00008212782,0.80226576,0.00011352537,0.000022156259,0.000088758425,0.00031805263,0.00029835996,0.0005798699],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99819934,0.0007064538,0.0001512917,0.00035668374,0.00046916213,0.000117022275],"domain_scores_gemma":[0.9873368,0.008777455,0.0011659621,0.0014404219,0.0011032424,0.00017623164],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030098283,0.0013430773,0.0007811107,0.0024871458,0.0005981526,0.0012350895,0.0019726257,0.0010038757,0.0020520398],"category_scores_gemma":[0.018812187,0.0006163286,0.00091395073,0.0011312271,0.0005716978,0.0024437471,0.000983657,0.001810542,0.0009714415],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005426982,0.0005464484,0.010416729,0.00038036972,0.00019251746,0.0006253872,0.000531788,0.36155495,0.0153779695,0.015865507,0.0050522042,0.5889134],"study_design_scores_gemma":[0.000021203163,0.00007451984,0.0004409619,0.0000265787,0.00002085235,0.00009577228,0.000016382126,0.9824724,0.007779044,0.0075968895,0.0014331695,0.000022234968],"about_ca_topic_score_codex":0.0035527723,"about_ca_topic_score_gemma":0.005096454,"teacher_disagreement_score":0.0035527723,"about_ca_system_score_codex":0.00066022214,"about_ca_system_score_gemma":0.0013295423,"threshold_uncertainty_score":0.015917659},"labels":[],"label_agreement":null},{"id":"W4232500090","doi":"10.32920/ryerson.14648862.v1","title":"Importance analysis of fault trees by visual inspection","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Fault tree analysis; Computer science; Generalization; Task (project management); Reliability (semiconductor); Rank (graph theory); Component (thermodynamics); Simple (philosophy); Tree (set theory); Fault (geology); Software; Reliability engineering; Data mining; Algorithm; Theoretical computer science; Machine learning; Mathematics; Engineering; Programming language","score_opus":0.014553449834480215,"score_gpt":0.301696943353354,"score_spread":0.2871434935188738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232500090","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020495016,0.000108721215,0.9777874,0.00004019395,0.0000127447665,0.000040689705,0.000060310475,0.0007326309,0.000722396],"genre_scores_gemma":[0.47332218,0.00023703965,0.52380854,0.00003487194,0.000042038802,0.000083174076,0.0003811927,0.00034537516,0.0017455629],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99872166,0.0003404948,0.000078487574,0.00019025092,0.00052439194,0.0001446144],"domain_scores_gemma":[0.9921449,0.004331698,0.0008590224,0.0006337932,0.0017612814,0.0002693238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015993036,0.0006480468,0.00066014414,0.0057246583,0.00042326574,0.001487824,0.0009870252,0.0006237223,0.002994666],"category_scores_gemma":[0.013618248,0.0004622454,0.00071148894,0.0017456416,0.0009411213,0.0016692161,0.0010549518,0.0009981428,0.00060868624],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059532514,0.0001449937,0.010493417,0.0005959562,0.00006808088,0.0005775413,0.00091240316,0.18508686,0.060572956,0.048564587,0.0042648525,0.68812305],"study_design_scores_gemma":[0.000025342753,0.00011462142,0.007364254,0.00004477613,0.000029742032,0.00033450313,0.00013441611,0.9124346,0.019813158,0.056592528,0.0030666254,0.00004546534],"about_ca_topic_score_codex":0.0021811966,"about_ca_topic_score_gemma":0.0015384736,"teacher_disagreement_score":0.0057246583,"about_ca_system_score_codex":0.00081644044,"about_ca_system_score_gemma":0.00047428635,"threshold_uncertainty_score":0.01001811},"labels":[],"label_agreement":null},{"id":"W4232560806","doi":"10.1109/msr.2015.73","title":"The Firefox Temporal Defect Dataset","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Process (computing); Software bug; Software; Plan (archaeology); Data science; Data mining; Geography","score_opus":0.04834508987867862,"score_gpt":0.3027732577530527,"score_spread":0.25442816787437406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232560806","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06378667,0.0011258124,0.0018186516,0.00049230823,0.00013089376,0.0002128609,0.92561257,0.0031821448,0.0036381497],"genre_scores_gemma":[0.01872324,0.00024959663,0.004790925,0.0001200639,0.000038886923,0.00018790246,0.97449625,0.00010546003,0.001287704],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988263,0.000105891384,0.00014481338,0.00033362923,0.0004417916,0.00014755105],"domain_scores_gemma":[0.9974011,0.0006411727,0.0004289594,0.0005143857,0.0007355566,0.00027881822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011416276,0.0015402018,0.00066423614,0.00624209,0.0007424774,0.0009105314,0.0019059669,0.0015364467,0.0029733188],"category_scores_gemma":[0.0043342523,0.0003995251,0.0010649967,0.0041435966,0.00036976868,0.0011062035,0.0009217242,0.0010535938,0.004137038],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008118515,0.0008798612,0.07959343,0.001692055,0.0002918326,0.0011890826,0.00040224113,0.0049408586,0.0053655636,0.0019289345,0.82379895,0.07910528],"study_design_scores_gemma":[0.0007905677,0.00070237304,0.2540245,0.0004610644,0.00027673822,0.0030974627,0.0006941147,0.024913078,0.0076301494,0.0031713392,0.70404077,0.0001978527],"about_ca_topic_score_codex":0.027402107,"about_ca_topic_score_gemma":0.053201176,"teacher_disagreement_score":0.027402107,"about_ca_system_score_codex":0.0014179697,"about_ca_system_score_gemma":0.0013863887,"threshold_uncertainty_score":0.0544852},"labels":[],"label_agreement":null},{"id":"W4232793935","doi":"10.1145/357474.355046","title":"Designing robust Java programs with exceptions","year":2000,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Exception handling; Java; Computer science; Structuring; Software engineering; Mechanism (biology); Source code; Programming language; Code (set theory); Set (abstract data type)","score_opus":0.02507598765568434,"score_gpt":0.23181477584984175,"score_spread":0.20673878819415742,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232793935","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027977452,0.00008973727,0.96579725,0.00016058584,0.000029011433,0.00017420344,0.000021355854,0.0038237178,0.0019267215],"genre_scores_gemma":[0.19516495,0.00033738645,0.7985139,0.00018046547,0.000046082045,0.00048019877,0.00015187648,0.0021310844,0.0029940298],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975114,0.0006537139,0.00027433143,0.00044239752,0.0008845609,0.00023357396],"domain_scores_gemma":[0.9953259,0.0020159318,0.00076956576,0.0009866654,0.00072613586,0.00017570377],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037302244,0.00061323866,0.00046241048,0.0005245369,0.0006791921,0.0019853134,0.0017998634,0.0011420806,0.0014287954],"category_scores_gemma":[0.0111593995,0.0011467962,0.0008177045,0.00033626411,0.0013357971,0.0024410216,0.0020705932,0.0016243553,0.0008277683],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000270295,0.00048830523,0.008987978,0.0013116527,0.00024948904,0.0020700623,0.0046227407,0.2794727,0.15888482,0.13023724,0.0075895316,0.4058153],"study_design_scores_gemma":[0.00028972994,0.00066567346,0.0017975139,0.00031077763,0.00037057235,0.0011237082,0.00066573184,0.716944,0.11555195,0.05744616,0.104676515,0.00015769222],"about_ca_topic_score_codex":0.0006357739,"about_ca_topic_score_gemma":0.0009018969,"teacher_disagreement_score":0.0037302244,"about_ca_system_score_codex":0.00039272633,"about_ca_system_score_gemma":0.0012642225,"threshold_uncertainty_score":0.019727528},"labels":[],"label_agreement":null},{"id":"W4233214837","doi":"10.1109/msr.2015.26","title":"Ecosystems in GitHub and a Method for Ecosystem Identification Using Reference Coupling","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ecosystem; Identification (biology); Computer science; Software; Isolation (microbiology); Data science; Coupling (piping); Software engineering; Ecology; Engineering","score_opus":0.1265117246568945,"score_gpt":0.36854390895167055,"score_spread":0.24203218429477605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233214837","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0402911,0.00026625,0.946442,0.00015344551,0.00005732505,0.00022337413,0.0006858866,0.006827593,0.0050530094],"genre_scores_gemma":[0.22363693,0.00022311242,0.76738083,0.00009014617,0.000028367238,0.0006076635,0.0021638176,0.0017744385,0.004094774],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9938612,0.002112707,0.0007232399,0.0011157257,0.0018393437,0.0003477924],"domain_scores_gemma":[0.98815536,0.0040538767,0.0014614485,0.003538259,0.0022456658,0.0005454008],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044907886,0.0010462397,0.00063087326,0.010424884,0.001309814,0.0034169916,0.001454371,0.0013169259,0.0031509823],"category_scores_gemma":[0.026770359,0.0010271339,0.0016616415,0.0076903724,0.0007950137,0.0048275176,0.0038232792,0.0014755384,0.0012227125],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004371615,0.00030646447,0.09063665,0.0011871869,0.0006212861,0.0024789853,0.011848162,0.039704602,0.02625839,0.124631286,0.013912715,0.687977],"study_design_scores_gemma":[0.00008877349,0.00030035264,0.07803223,0.000528629,0.00046317265,0.0034086648,0.0036404782,0.5822836,0.027999377,0.14455442,0.15829425,0.0004059526],"about_ca_topic_score_codex":0.008084127,"about_ca_topic_score_gemma":0.010967918,"teacher_disagreement_score":0.010424884,"about_ca_system_score_codex":0.0011242052,"about_ca_system_score_gemma":0.0016994495,"threshold_uncertainty_score":0.023749888},"labels":[],"label_agreement":null},{"id":"W4233337392","doi":"10.1145/1083125.1083131","title":"ActiveAspect","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"AspectJ; Computer science; Sketch; Presentation (obstetrics); Abstraction; Software engineering; Programming language; Data structure; Aspect-oriented programming; Software; Algorithm","score_opus":0.014851547825062856,"score_gpt":0.2679253769666525,"score_spread":0.2530738291415896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233337392","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008259143,0.00065802335,0.68924004,0.0009840039,0.0004584852,0.00042205563,0.0026035404,0.24281523,0.05455949],"genre_scores_gemma":[0.13580763,0.0016638904,0.62514997,0.0025409048,0.00043431035,0.0010882503,0.017720154,0.09246668,0.12312821],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99808145,0.000348334,0.00019265141,0.00027005564,0.0009771448,0.00013033445],"domain_scores_gemma":[0.99459654,0.0022997754,0.0003367517,0.0013704136,0.0010232303,0.00037337668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027277195,0.0014080013,0.0006312483,0.0012885517,0.00092003087,0.003775033,0.0031124952,0.0025324444,0.027661053],"category_scores_gemma":[0.01052134,0.0015293981,0.0013463597,0.000613583,0.0008242622,0.0071804975,0.004485358,0.003555552,0.017518107],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009948083,0.00040378803,0.0056031505,0.0015872109,0.00015826717,0.001495609,0.0037801326,0.0023018813,0.02962276,0.13914828,0.34922642,0.46567765],"study_design_scores_gemma":[0.00009131933,0.00011881787,0.0013686167,0.000176482,0.000060291022,0.0016004322,0.00022185466,0.011576848,0.018688157,0.03475668,0.9312588,0.00008177128],"about_ca_topic_score_codex":0.0011322038,"about_ca_topic_score_gemma":0.0018412495,"teacher_disagreement_score":0.027661053,"about_ca_system_score_codex":0.0005739535,"about_ca_system_score_gemma":0.001061585,"threshold_uncertainty_score":0.092535436},"labels":[],"label_agreement":null},{"id":"W4233384578","doi":"10.1109/qsic.2004.1357950","title":"Machine-learning techniques for software product quality assessment","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Computer science; Software quality; Predictability; Quality (philosophy); Software measurement; Verification and validation; Software construction; Software sizing; Software metric; Software; Software engineering; Software quality control; Software development; Machine learning; Artificial intelligence; Engineering","score_opus":0.04024851804344539,"score_gpt":0.3565713603113449,"score_spread":0.31632284226789953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233384578","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031409631,0.0036230942,0.9908647,0.0002997764,0.000049467002,0.000051465413,0.0000799841,0.0004810456,0.0014094888],"genre_scores_gemma":[0.28202695,0.0071353144,0.70555973,0.00023168851,0.0004242949,0.0004886089,0.0006516001,0.00010573486,0.0033760848],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980731,0.00076406443,0.00017530027,0.00023717538,0.00068999623,0.000060421055],"domain_scores_gemma":[0.99549955,0.0033402278,0.00030797,0.000300638,0.0005168354,0.00003479427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024317445,0.0009952228,0.0010503093,0.0025493526,0.00039776613,0.0013112251,0.0014659788,0.0012768054,0.0020143536],"category_scores_gemma":[0.010513394,0.00034596043,0.00077698834,0.0037570747,0.0007178858,0.0017359076,0.0008834852,0.0016560717,0.0012168938],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004822988,0.00011814919,0.0017349516,0.0004447802,0.00011409183,0.00007681328,0.000088185196,0.32911983,0.0014325914,0.032427296,0.00306775,0.63132733],"study_design_scores_gemma":[0.000008413135,0.000033398363,0.0005297791,0.000046394027,0.000015996384,0.000048119655,0.000016713497,0.94583553,0.0009442581,0.04916559,0.0033403477,0.000015524172],"about_ca_topic_score_codex":0.002627623,"about_ca_topic_score_gemma":0.001921333,"teacher_disagreement_score":0.002627623,"about_ca_system_score_codex":0.00086400745,"about_ca_system_score_gemma":0.0007918514,"threshold_uncertainty_score":0.012860477},"labels":[],"label_agreement":null},{"id":"W4233479913","doi":"10.1109/msr.2012.6224280","title":"Explaining software defects using topic models","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Software quality; Software metric; Eclipse; Compiler; Code review; Software engineering; Software development; Source lines of code; Software quality assurance; Static program analysis; Software; Code smell; Data science; Domain (mathematical analysis); Software system; Source code; Software bug; Quality (philosophy); Programming language","score_opus":0.06928329702945245,"score_gpt":0.29332141752629654,"score_spread":0.22403812049684407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233479913","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20372969,0.0022004205,0.78602815,0.0017252836,0.00017077425,0.0003701217,0.0018636909,0.0017305963,0.00218125],"genre_scores_gemma":[0.8926523,0.0013863468,0.09932383,0.00018653947,0.00034371854,0.0007823244,0.0031305107,0.0003009947,0.0018934506],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.99625,0.0019029888,0.00026804203,0.0009126113,0.00043314177,0.00023327193],"domain_scores_gemma":[0.9445788,0.048994437,0.002892058,0.0013232677,0.0017836748,0.00042776405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009202563,0.0019796262,0.0016792878,0.0091625415,0.0010468494,0.0033753272,0.0024535449,0.0025237685,0.0025412827],"category_scores_gemma":[0.038415305,0.00093881745,0.0036271957,0.005656901,0.0011516891,0.005300748,0.0017460819,0.0020838303,0.0008462512],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007689425,0.00063531724,0.12368029,0.0010045979,0.0013288544,0.0007501006,0.0063866028,0.49428433,0.003987483,0.08440358,0.011718387,0.27105144],"study_design_scores_gemma":[0.00004173012,0.00006208135,0.0066020316,0.00005028091,0.00013187445,0.00013570655,0.0002406235,0.9578302,0.00033655734,0.033026885,0.0015023001,0.00003982022],"about_ca_topic_score_codex":0.011511008,"about_ca_topic_score_gemma":0.007998476,"teacher_disagreement_score":0.011511008,"about_ca_system_score_codex":0.0020174186,"about_ca_system_score_gemma":0.0010471322,"threshold_uncertainty_score":0.048668444},"labels":[],"label_agreement":null},{"id":"W4233947233","doi":"10.1109/cesi.2015.10","title":"Planning for the Unknown: Lessons Learned from Ten Months of Non-participant Exploratory Observations in the Industry","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Exploratory research; Observational study; Process (computing); Knowledge management; Computer science; Software; Management science; Process management; Engineering; Medicine; Sociology","score_opus":0.4450473984073209,"score_gpt":0.39483790261366436,"score_spread":0.05020949579365652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233947233","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88153297,0.0028237244,0.075332485,0.019992065,0.00066583103,0.0027366986,0.00045846702,0.0006097113,0.015847903],"genre_scores_gemma":[0.9588445,0.0014753749,0.031856522,0.0021871785,0.0001462579,0.0016622744,0.00025605175,0.00024132623,0.0033305134],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9354846,0.05221539,0.0016599903,0.0032695457,0.004349971,0.0030205382],"domain_scores_gemma":[0.7338389,0.205357,0.007936653,0.01699132,0.024692604,0.01118347],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08854695,0.001329246,0.0015896965,0.0016898201,0.011017951,0.0073127397,0.0059701456,0.0037951209,0.0033593648],"category_scores_gemma":[0.17688753,0.00122044,0.00088015496,0.0015684217,0.009349714,0.013258608,0.009506038,0.0072358013,0.0013806006],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042935542,0.0017938166,0.022484498,0.001024403,0.00006912298,0.0025858725,0.8564603,0.0005786398,0.0018810667,0.0018516243,0.0072218324,0.1036195],"study_design_scores_gemma":[0.00008886259,0.0010129673,0.023009034,0.001674924,0.000063308886,0.0008376114,0.92197573,0.0013305087,0.0014238857,0.0069192797,0.04147217,0.00019172294],"about_ca_topic_score_codex":0.011074312,"about_ca_topic_score_gemma":0.028933978,"teacher_disagreement_score":0.91145307,"about_ca_system_score_codex":0.004995027,"about_ca_system_score_gemma":0.008721334,"threshold_uncertainty_score":0.46828663},"labels":[],"label_agreement":null},{"id":"W4233966475","doi":"10.1109/weh.2012.6226604","title":"Usability challenges in exception handling","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Exception handling; Assertion; Usability; Matching (statistics); Programming language; Block (permutation group theory); Human–computer interaction","score_opus":0.09250696725919971,"score_gpt":0.30998267618308695,"score_spread":0.21747570892388723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233966475","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16887023,0.007965681,0.71456707,0.039986618,0.001554208,0.0015478315,0.00027564212,0.017096547,0.048136175],"genre_scores_gemma":[0.53058696,0.0030875248,0.43757072,0.0060080728,0.00086680334,0.00075090484,0.00038163437,0.0063823685,0.014365032],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8546676,0.08629761,0.011245122,0.007351369,0.037167512,0.0032708074],"domain_scores_gemma":[0.5442146,0.31233153,0.012777992,0.057765223,0.06918514,0.0037255574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10867852,0.0014469415,0.0014217177,0.0036426852,0.0033922214,0.017914945,0.006771242,0.0040820907,0.005594454],"category_scores_gemma":[0.27455375,0.0014096717,0.001653318,0.002485113,0.0064251553,0.017007848,0.00720362,0.005084015,0.0026400439],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010325505,0.0005931404,0.018863019,0.00531997,0.00021953104,0.0026529701,0.06009917,0.003758725,0.021424739,0.0836159,0.028085196,0.7743351],"study_design_scores_gemma":[0.0004720915,0.003678015,0.015026778,0.01042371,0.000828672,0.015493914,0.06483206,0.05244597,0.033540953,0.30530727,0.49668103,0.001269559],"about_ca_topic_score_codex":0.0050204378,"about_ca_topic_score_gemma":0.004451597,"teacher_disagreement_score":0.10867852,"about_ca_system_score_codex":0.0032042435,"about_ca_system_score_gemma":0.00580226,"threshold_uncertainty_score":0.5747538},"labels":[],"label_agreement":null},{"id":"W4234320625","doi":"10.1109/msr.2013.6623994","title":"Contents","year":2013,"lang":"en","type":"paratext","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Saskatchewan; University of Waterloo; University of Toronto","funders":"","keywords":"Computer science","score_opus":0.03462356911873928,"score_gpt":0.286197415755017,"score_spread":0.25157384663627774,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4234320625","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022714694,0.0060457205,0.016249996,0.00838977,0.01350839,0.0005989603,0.033179812,0.005305513,0.91445035],"genre_scores_gemma":[0.005915695,0.005028173,0.005031005,0.0015595147,0.005956452,0.0003088764,0.028857293,0.002012949,0.94533],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993585,0.00010819911,0.000038985978,0.00011466018,0.00032916942,0.000050521034],"domain_scores_gemma":[0.99658954,0.0007719822,0.00017187657,0.0004725617,0.0013408652,0.0006532213],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0008179921,0.0007612801,0.0008263033,0.005062769,0.0013243373,0.0065527824,0.0011637117,0.00086293457,0.6364426],"category_scores_gemma":[0.0065201577,0.00031798045,0.00046072007,0.0062151486,0.00048524854,0.0042589568,0.0027221658,0.0012245044,0.52659184],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017246048,0.00004050343,0.00022638234,0.0001588713,0.000002697656,0.000018032435,0.000047045065,0.00016643743,0.00034769138,0.005043739,0.9070303,0.08690101],"study_design_scores_gemma":[0.0000036663341,0.000013533023,0.0005052395,0.00015788298,0.0000030129293,0.00004821013,0.000059727077,0.0002074927,0.0001655114,0.0028980323,0.99593335,0.0000043913597],"about_ca_topic_score_codex":0.0017180631,"about_ca_topic_score_gemma":0.002007007,"teacher_disagreement_score":0.3635574,"about_ca_system_score_codex":0.0017539882,"about_ca_system_score_gemma":0.0018390992,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4234560745","doi":"10.1145/512058.512062","title":"Tracking structural evolution using origin analysis","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Software; Term (time); Software system; Software architecture; Software engineering; Work (physics); Software construction; Engineering; Programming language","score_opus":0.054518655594594144,"score_gpt":0.30327786483544283,"score_spread":0.24875920924084868,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4234560745","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.095500834,0.00044936204,0.89705354,0.00016400457,0.00003663973,0.000050032828,0.00033382906,0.0040786276,0.0023331076],"genre_scores_gemma":[0.61773604,0.00043227055,0.37813762,0.000046809153,0.00002861015,0.0000987836,0.00090572244,0.00054803136,0.0020660397],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99887985,0.00023043978,0.000053303203,0.00033015426,0.00042603153,0.00008018359],"domain_scores_gemma":[0.9946642,0.0020279663,0.0010996322,0.0011835129,0.00089842366,0.00012622956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014369156,0.0005267072,0.0005613947,0.0033508854,0.00073836977,0.0012943539,0.0012970861,0.0010811945,0.0016100039],"category_scores_gemma":[0.009409913,0.0005072861,0.0006043631,0.0019914242,0.0009748483,0.0024193826,0.0020801262,0.0013353754,0.00047117597],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056297006,0.00019935709,0.09957417,0.00048087438,0.00023591661,0.0011159562,0.0033902975,0.19284266,0.07703513,0.10020293,0.0041502165,0.52020943],"study_design_scores_gemma":[0.000029192466,0.00008413169,0.009226066,0.000034192162,0.000078754514,0.0003188486,0.0001824263,0.9228682,0.02241819,0.035676263,0.009039912,0.0000437533],"about_ca_topic_score_codex":0.0033806781,"about_ca_topic_score_gemma":0.0023624047,"teacher_disagreement_score":0.0033806781,"about_ca_system_score_codex":0.0007626332,"about_ca_system_score_gemma":0.0005566937,"threshold_uncertainty_score":0.0075992346},"labels":[],"label_agreement":null},{"id":"W4234685938","doi":"10.32920/ryerson.14648862","title":"Importance analysis of fault trees by visual inspection","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Fault tree analysis; Computer science; Generalization; Reliability (semiconductor); Task (project management); Rank (graph theory); Component (thermodynamics); Simple (philosophy); Reliability engineering; Fault (geology); Tree (set theory); Software; Algorithm; Data mining; Theoretical computer science; Mathematics; Programming language; Engineering","score_opus":0.014553449834480215,"score_gpt":0.301696943353354,"score_spread":0.2871434935188738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4234685938","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020495016,0.000108721215,0.9777874,0.00004019395,0.0000127447665,0.000040689705,0.000060310475,0.0007326309,0.000722396],"genre_scores_gemma":[0.47332218,0.00023703965,0.52380854,0.00003487194,0.000042038802,0.000083174076,0.0003811927,0.00034537516,0.0017455629],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99872166,0.0003404948,0.000078487574,0.00019025092,0.00052439194,0.0001446144],"domain_scores_gemma":[0.9921449,0.004331698,0.0008590224,0.0006337932,0.0017612814,0.0002693238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015993036,0.0006480468,0.00066014414,0.0057246583,0.00042326574,0.001487824,0.0009870252,0.0006237223,0.002994666],"category_scores_gemma":[0.013618248,0.0004622454,0.00071148894,0.0017456416,0.0009411213,0.0016692161,0.0010549518,0.0009981428,0.00060868624],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059532514,0.0001449937,0.010493417,0.0005959562,0.00006808088,0.0005775413,0.00091240316,0.18508686,0.060572956,0.048564587,0.0042648525,0.68812305],"study_design_scores_gemma":[0.000025342753,0.00011462142,0.007364254,0.00004477613,0.000029742032,0.00033450313,0.00013441611,0.9124346,0.019813158,0.056592528,0.0030666254,0.00004546534],"about_ca_topic_score_codex":0.0021811966,"about_ca_topic_score_gemma":0.0015384736,"teacher_disagreement_score":0.0057246583,"about_ca_system_score_codex":0.00081644044,"about_ca_system_score_gemma":0.00047428635,"threshold_uncertainty_score":0.01001811},"labels":[],"label_agreement":null},{"id":"W4234703843","doi":"10.1109/msr.2015.34","title":"An Empirical Study of End-User Programmers in the Computer Music Community","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Metadata; Software; World Wide Web; Software development; Computer programming; Population; Source code; End user; Musical; Empirical research; Software engineering; Programming language","score_opus":0.13095042248465816,"score_gpt":0.3737612601332305,"score_spread":0.24281083764857236,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4234703843","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9990159,0.00005005887,0.00012971721,0.00012429111,0.0000020919938,0.000022866421,0.00001707547,0.0000030630401,0.00063478114],"genre_scores_gemma":[0.9984688,0.00016731572,0.00042873647,0.00017657674,0.000008838415,0.000057568886,0.000070012,0.0000075020384,0.0006146851],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9929382,0.0035521823,0.00040881173,0.00066255993,0.0016076164,0.0008305808],"domain_scores_gemma":[0.91344506,0.048816916,0.014354268,0.001911377,0.011441077,0.010031277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064020464,0.00029305834,0.00039480985,0.00398307,0.0032541628,0.0027822808,0.0012809547,0.001065929,0.002399293],"category_scores_gemma":[0.04807297,0.0005016991,0.00013943556,0.003285988,0.0022827783,0.003738081,0.0027983645,0.0016584062,0.00037458292],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014175363,0.0012744134,0.58257973,0.00018368686,0.000022169661,0.0011172714,0.3890081,0.00005012848,0.0012266302,0.00058363425,0.0010736192,0.022738844],"study_design_scores_gemma":[0.000026414134,0.0005820408,0.29032606,0.0001182079,0.000012988993,0.0009841307,0.70238346,0.0005170929,0.0003992008,0.0002607265,0.004354463,0.000035197943],"about_ca_topic_score_codex":0.006249608,"about_ca_topic_score_gemma":0.012320883,"teacher_disagreement_score":0.0064020464,"about_ca_system_score_codex":0.00097295764,"about_ca_system_score_gemma":0.0020157776,"threshold_uncertainty_score":0.033857644},"labels":[],"label_agreement":null},{"id":"W4234750371","doi":"10.4018/9781591409411.ch008","title":"Modeling Relevance Relations Using Machine Learning Techniques","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Relevance (law); Relation (database); Artificial intelligence; Machine learning; Software; Software deployment; Abstraction; Precision and recall; Tuple; Classifier (UML); Data mining; Information retrieval; Data science; Software engineering; Programming language","score_opus":0.04127005045233979,"score_gpt":0.27500254848819566,"score_spread":0.23373249803585588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4234750371","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024038525,0.0034776195,0.96192294,0.0015093792,0.00010859195,0.00025020883,0.0006701801,0.0012795704,0.0067429077],"genre_scores_gemma":[0.41753078,0.0030247003,0.56881976,0.00051029574,0.00044128063,0.0006333013,0.0023536761,0.0003599689,0.0063262153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937289,0.002356597,0.0004315272,0.0012635421,0.0018900317,0.00032931744],"domain_scores_gemma":[0.9791027,0.017297586,0.0011903353,0.0010273581,0.0012067292,0.00017533886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059495433,0.001417273,0.001514208,0.0067682574,0.0013945411,0.0037175966,0.0026358627,0.001982934,0.004241503],"category_scores_gemma":[0.03148562,0.0009906853,0.0020706465,0.0051889718,0.0013389545,0.0077887233,0.002034723,0.0032247764,0.0016743571],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019689974,0.00043816754,0.0135004325,0.0007046642,0.0003054164,0.0008802291,0.0012525,0.3963636,0.0022949637,0.14696962,0.017218135,0.4198754],"study_design_scores_gemma":[0.000015048234,0.000029172515,0.0011540798,0.000078997575,0.000040183542,0.00018432798,0.0000834562,0.85266596,0.00062804,0.13883781,0.0062589864,0.000023917353],"about_ca_topic_score_codex":0.007233726,"about_ca_topic_score_gemma":0.0077275927,"teacher_disagreement_score":0.007233726,"about_ca_system_score_codex":0.0023671484,"about_ca_system_score_gemma":0.0016945815,"threshold_uncertainty_score":0.031464577},"labels":[],"label_agreement":null},{"id":"W4235705087","doi":"10.32920/ryerson.14648622","title":"On Empirically Examining The Effectiveness Of Deep Learning-Based Bug Localization Models","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software bug; Java; Set (abstract data type); Software; Artificial intelligence; Convolutional neural network; Deep learning; Baseline (sea); Convolution (computer science); Machine learning; Software engineering; State (computer science); Artificial neural network; Programming language","score_opus":0.033230855442043704,"score_gpt":0.2871205612570827,"score_spread":0.253889705815039,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4235705087","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9066384,0.008643059,0.0666111,0.003691472,0.0003315026,0.00025496355,0.0021807505,0.0021593135,0.009489511],"genre_scores_gemma":[0.97249776,0.001058041,0.02262434,0.00037232458,0.00006445728,0.00009948421,0.0022163386,0.00009136518,0.00097589183],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99542266,0.0021925499,0.0004475129,0.0008661835,0.0007376974,0.00033332955],"domain_scores_gemma":[0.9391235,0.04854934,0.0029101355,0.003895964,0.004634511,0.00088658824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012441448,0.0031701184,0.0013602679,0.0020496568,0.00084451283,0.0021195032,0.0022226898,0.003560991,0.0018640254],"category_scores_gemma":[0.060674608,0.00088976504,0.0009982137,0.0018733937,0.0015750283,0.006096322,0.0017745451,0.0036720482,0.0007127594],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014326922,0.0012363022,0.04046419,0.0008767899,0.0006226913,0.00014404485,0.00018573357,0.84326476,0.0022521361,0.0031286576,0.0070707765,0.09932121],"study_design_scores_gemma":[0.00008840155,0.00067299936,0.0035010257,0.000146479,0.00014844288,0.00006563409,0.000096204734,0.9891777,0.0024119809,0.002928714,0.00073075824,0.000031628108],"about_ca_topic_score_codex":0.021703608,"about_ca_topic_score_gemma":0.021387901,"teacher_disagreement_score":0.021703608,"about_ca_system_score_codex":0.0031580415,"about_ca_system_score_gemma":0.0018400372,"threshold_uncertainty_score":0.06579739},"labels":[],"label_agreement":null},{"id":"W4235970556","doi":"10.1007/s10664-021-10009-1","title":"Conclusion stability for natural language based mining of design discussions","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Artifact (error); Code refactoring; Computer science; Relevance (law); Task (project management); Context (archaeology); Documentation; Stability (learning theory); Software; Machine learning; Artificial intelligence; Data mining; Software engineering; Engineering; Programming language; Systems engineering","score_opus":0.04044770805115119,"score_gpt":0.3104229520862787,"score_spread":0.2699752440351275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4235970556","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.114228204,0.0006064289,0.8683203,0.0015344034,0.00013686516,0.0003135786,0.0048999796,0.0043619056,0.0055983276],"genre_scores_gemma":[0.81178266,0.00024903897,0.17403032,0.0004074845,0.00018219544,0.00046872115,0.007869638,0.0007573651,0.0042525823],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98853904,0.0031442377,0.0010264288,0.0034957202,0.0031161695,0.0006784639],"domain_scores_gemma":[0.9044801,0.06735961,0.004227835,0.010062193,0.012591362,0.0012788267],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009759915,0.0008006552,0.0012231056,0.0061818943,0.002081123,0.0045421235,0.0022697428,0.0013851052,0.010729092],"category_scores_gemma":[0.094172716,0.00063927524,0.0021078405,0.0030748595,0.0019244993,0.0072639263,0.0033878458,0.002233748,0.003211234],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033898847,0.0008010621,0.08377755,0.0017141671,0.00053238875,0.0007561265,0.0048597464,0.08771267,0.0347743,0.24976005,0.029687168,0.5022349],"study_design_scores_gemma":[0.00009342159,0.00024388322,0.009610051,0.0001787125,0.0001502574,0.0003690852,0.0011551714,0.7075464,0.023066888,0.24840015,0.00910558,0.00008050787],"about_ca_topic_score_codex":0.004450572,"about_ca_topic_score_gemma":0.003429539,"teacher_disagreement_score":0.9902401,"about_ca_system_score_codex":0.002320975,"about_ca_system_score_gemma":0.002666763,"threshold_uncertainty_score":0.051616013},"labels":[],"label_agreement":null},{"id":"W4236048929","doi":"10.1002/smr.376","title":"Introduction to the special issue on program comprehension through dynamic analysis (PCODA)","year":2008,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Concordia University","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung","keywords":"Program comprehension; Computer science; Reverse engineering; Comprehension; Software engineering; Software; Program analysis; Source code; Data science; Software system; Programming language","score_opus":0.03584303218008123,"score_gpt":0.3541619928238779,"score_spread":0.31831896064379667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236048929","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005736085,0.035446752,0.022870392,0.0463606,0.8431569,0.00033429617,0.001353294,0.0025907059,0.047313455],"genre_scores_gemma":[0.0034944764,0.03537858,0.008589917,0.022132268,0.77694005,0.0005290898,0.0036431393,0.00538775,0.1439047],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99407876,0.0009637527,0.00062562525,0.0013520725,0.0024844133,0.00049546186],"domain_scores_gemma":[0.97554207,0.010348525,0.0013310768,0.0019932098,0.0071759424,0.0036092403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047209114,0.0025149754,0.0027845546,0.00494455,0.0024536685,0.011824849,0.0029452988,0.0043147597,0.1325066],"category_scores_gemma":[0.019128732,0.0010785265,0.0027075345,0.0038996227,0.0021206955,0.009460911,0.0059721298,0.009294904,0.06895125],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003105224,0.000040929037,0.00008392569,0.00039150845,0.000013870572,0.000049208946,0.00006844937,0.00012148196,0.0003881982,0.0021693264,0.9513757,0.045266308],"study_design_scores_gemma":[0.000007826404,0.000028342936,0.00018916666,0.00029308934,0.0000082604265,0.00012296063,0.000039076713,0.00014930959,0.00013119985,0.0024391047,0.9965771,0.000014517811],"about_ca_topic_score_codex":0.0010268693,"about_ca_topic_score_gemma":0.0013518197,"teacher_disagreement_score":0.1325066,"about_ca_system_score_codex":0.002821837,"about_ca_system_score_gemma":0.003394821,"threshold_uncertainty_score":0.44327873},"labels":[],"label_agreement":null},{"id":"W4236085600","doi":"10.1109/msr.2015.47","title":"Organizational Volatility and Post-release Defects: A Replication Case Study Using Data from Google Chrome","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Outsourcing; Popularity; Volatility (finance); Business; Computer science; Directory; World Wide Web; Replicate; Knowledge management; Marketing; Finance; Operating system","score_opus":0.09185613977482662,"score_gpt":0.33546899444422806,"score_spread":0.24361285466940144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236085600","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9984794,0.00005399254,0.0005624961,0.00006199675,0.000004810715,0.00008127718,0.00041879268,0.000021570308,0.00031565572],"genre_scores_gemma":[0.9971308,0.000040630774,0.0014074622,0.000038655522,0.000013253412,0.00012908432,0.00090210006,0.000018438835,0.0003196806],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9915073,0.004553551,0.000710334,0.0012793598,0.0014291345,0.00052030367],"domain_scores_gemma":[0.9110383,0.048645984,0.013946695,0.014264087,0.010406834,0.0016981034],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009701415,0.00053937803,0.0006198673,0.0030284596,0.0013413727,0.0014376445,0.0019871632,0.0017210748,0.0010748841],"category_scores_gemma":[0.046714276,0.00047671422,0.0010949814,0.0034458141,0.0013332919,0.0015593902,0.001355434,0.0014482243,0.0004588512],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006773879,0.0017929365,0.9534969,0.00024610228,0.0003370039,0.0036891026,0.01689617,0.002248683,0.0021066018,0.0005453019,0.0017163472,0.016247492],"study_design_scores_gemma":[0.0001893666,0.0015411079,0.9600187,0.00008831395,0.00021744828,0.001709106,0.02090094,0.008436034,0.0026884011,0.00053515815,0.003540094,0.00013526596],"about_ca_topic_score_codex":0.02855391,"about_ca_topic_score_gemma":0.0281674,"teacher_disagreement_score":0.99029857,"about_ca_system_score_codex":0.0015429043,"about_ca_system_score_gemma":0.0011259093,"threshold_uncertainty_score":0.05677539},"labels":[],"label_agreement":null},{"id":"W4236348706","doi":"10.1002/spip.433","title":"Strengthening maturity levels by a legal assurance process","year":2009,"lang":"en","type":"article","venue":"Software Process Improvement and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Capability Maturity Model Integration; Maturity (psychological); Process (computing); Process management; Capability Maturity Model; Engineering; Business; Engineering management; Computer science; Software development process; Law; Software development; Political science; Software","score_opus":0.016295764161487238,"score_gpt":0.31201564117518554,"score_spread":0.2957198770136983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236348706","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16200386,0.002072995,0.7178097,0.014663102,0.00025682285,0.0015721658,0.00014557764,0.0022502558,0.09922563],"genre_scores_gemma":[0.6225624,0.0004870703,0.36983246,0.00042672615,0.00010381081,0.0007024475,0.0001666354,0.00013019526,0.005588234],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96835303,0.016679205,0.0026229052,0.001677092,0.009275461,0.0013922842],"domain_scores_gemma":[0.9178941,0.032027658,0.011616664,0.0076030586,0.026568346,0.0042901887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.033537034,0.0006669133,0.00034097844,0.0043586204,0.0021934605,0.009619146,0.001513336,0.0024169935,0.0038046387],"category_scores_gemma":[0.07555661,0.00056844205,0.00063851225,0.00302661,0.0026650538,0.012065812,0.0064391484,0.0029095793,0.0014840942],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015795436,0.0005303253,0.015424963,0.00079268747,0.000050898245,0.00034335037,0.015971398,0.0063903104,0.01188031,0.39978355,0.0061971936,0.5424771],"study_design_scores_gemma":[0.00023261573,0.00189693,0.03065298,0.0031798547,0.00018972674,0.0007537792,0.013548504,0.079075605,0.032018855,0.5191731,0.31890646,0.00037161255],"about_ca_topic_score_codex":0.00185158,"about_ca_topic_score_gemma":0.0012801822,"teacher_disagreement_score":0.033537034,"about_ca_system_score_codex":0.006098848,"about_ca_system_score_gemma":0.0114791,"threshold_uncertainty_score":0.17736286},"labels":[],"label_agreement":null},{"id":"W4236372141","doi":"10.14393/ufu.te.2020.454","title":"CROKAGE: effective solution recommendation for programming tasks by leveraging crowd knowledge","year":2020,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Saskatchewan Health","funders":"","keywords":"Computer science; Leverage (statistics); Information retrieval; Code (set theory); Task (project management); Relevance (law); Crowdsourcing; Programming language; World Wide Web; Artificial intelligence","score_opus":0.021094536520200356,"score_gpt":0.3185644232366015,"score_spread":0.2974698867164011,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236372141","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.104657,0.006311343,0.7730718,0.0024109604,0.00055771205,0.0028865822,0.0074006803,0.08621307,0.016490819],"genre_scores_gemma":[0.20801573,0.0011857006,0.7649976,0.0011076921,0.00013961691,0.0011478995,0.01491482,0.0012811114,0.0072098193],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969228,0.00069125276,0.00019516675,0.0009793523,0.0010026869,0.00020880527],"domain_scores_gemma":[0.9969369,0.0015226441,0.0002475767,0.00048646465,0.0005564329,0.00025002635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024799483,0.003824936,0.0018251869,0.007657311,0.0013004539,0.0020410935,0.0036855408,0.0032260334,0.0064608846],"category_scores_gemma":[0.011547129,0.0008858762,0.0018329996,0.0029505375,0.000940808,0.00445934,0.0037424362,0.0021814145,0.0038040462],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012600466,0.0016483786,0.006646848,0.0024863726,0.00037055372,0.00093211944,0.0017140687,0.03931752,0.024571337,0.0066718245,0.11241331,0.8019675],"study_design_scores_gemma":[0.0006047831,0.00064276997,0.0033841666,0.00026703643,0.0002795717,0.00053518155,0.0017557846,0.8796058,0.01966091,0.02563729,0.06740096,0.00022575358],"about_ca_topic_score_codex":0.017086655,"about_ca_topic_score_gemma":0.033857957,"teacher_disagreement_score":0.017086655,"about_ca_system_score_codex":0.0013772538,"about_ca_system_score_gemma":0.00327793,"threshold_uncertainty_score":0.03397441},"labels":[],"label_agreement":null},{"id":"W4236561953","doi":"10.1145/1463788","title":"Proceedings of the 2008 conference of the center for advanced studies on collaborative research meeting of minds - CASCON '08","year":2008,"lang":"en","type":"paratext","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"IBM; Session (web analytics); Library science; Government (linguistics); Variety (cybernetics); Computer science; Research center; Publishing; Political science; World Wide Web","score_opus":0.106185880035104,"score_gpt":0.3811388353502063,"score_spread":0.27495295531510233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236561953","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01780778,0.043294754,0.024414033,0.21212682,0.2602479,0.0013294837,0.0046758032,0.0037340878,0.4323694],"genre_scores_gemma":[0.02740744,0.010968004,0.010015653,0.009546978,0.0124480445,0.0010375206,0.0052111656,0.0016247461,0.92174035],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9936818,0.0021038777,0.0003009795,0.000881581,0.0023067936,0.0007249883],"domain_scores_gemma":[0.9825602,0.0017721832,0.0004062125,0.0010100341,0.0060530882,0.008198335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012450204,0.0011188058,0.0013750264,0.0012522218,0.0034382525,0.01150822,0.0022495,0.002548173,0.1817904],"category_scores_gemma":[0.013973045,0.0004756974,0.0007710419,0.0013866145,0.0015574812,0.0051015727,0.007117128,0.0045328597,0.07554928],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000094087365,0.00007803189,0.00024085309,0.00014001112,0.000006891667,0.000047617625,0.0004454794,0.00005129858,0.00032692816,0.0031666637,0.9604544,0.034947723],"study_design_scores_gemma":[0.000014877487,0.000034214943,0.0007547386,0.00011640043,0.000004447639,0.000034146,0.00045206366,0.00009984896,0.00009212353,0.00086551916,0.99752,0.000011514404],"about_ca_topic_score_codex":0.0029130646,"about_ca_topic_score_gemma":0.008430071,"teacher_disagreement_score":0.1817904,"about_ca_system_score_codex":0.0036323695,"about_ca_system_score_gemma":0.006751155,"threshold_uncertainty_score":0.60814947},"labels":[],"label_agreement":null},{"id":"W4236722288","doi":"10.1145/1269899.1254904","title":"Synthetic designs","year":2007,"lang":"en","type":"article","venue":"ACM SIGMETRICS Performance Evaluation Review","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Communications Research Centre Canada","funders":"","keywords":"Computer science; Design of experiments; Reliability engineering; Statistical power; Variation (astronomy); Sample size determination; Sample (material); Statistical hypothesis testing; Software; Industrial engineering; Engineering; Statistics; Mathematics","score_opus":0.12903920978461722,"score_gpt":0.380485539652118,"score_spread":0.2514463298675008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236722288","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04162635,0.002602504,0.84679717,0.0028600027,0.0045738947,0.039983634,0.006912773,0.0011115194,0.05353221],"genre_scores_gemma":[0.22884597,0.0020498703,0.5981674,0.004141832,0.00092459016,0.13720423,0.0037537203,0.0003193056,0.024592998],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94815105,0.04016905,0.0029482243,0.0034493802,0.004449717,0.00083261606],"domain_scores_gemma":[0.8969189,0.06810164,0.006962292,0.015430265,0.011173379,0.0014135902],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.043597817,0.0015370174,0.0014297476,0.0016295399,0.0013160155,0.0026793617,0.0027271905,0.0022359819,0.051093128],"category_scores_gemma":[0.1066777,0.0008066288,0.0017116577,0.0014596615,0.0021445765,0.0024785898,0.0031753697,0.0027178945,0.0064305244],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013045618,0.0024511241,0.007846197,0.005393896,0.0010450399,0.000356884,0.0023118628,0.02066133,0.0036145537,0.4769865,0.054376602,0.4119105],"study_design_scores_gemma":[0.009305533,0.019375676,0.0055950144,0.0022093423,0.0010032534,0.0007786725,0.001620203,0.04714981,0.006756343,0.48695758,0.41889593,0.00035259008],"about_ca_topic_score_codex":0.00069063297,"about_ca_topic_score_gemma":0.00096774223,"teacher_disagreement_score":0.051093128,"about_ca_system_score_codex":0.0020745026,"about_ca_system_score_gemma":0.003374455,"threshold_uncertainty_score":0.23057008},"labels":[],"label_agreement":null},{"id":"W4236803832","doi":"10.17771/pucrio.acad.24701","title":"DETECÇÃO DE ANOMALIAS DE CÓDIGO DE RELEVÂNCIA ARQUITETURAL EM SISTEMAS MULTILINGUAGEM","year":2014,"lang":"pt","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Future Earth","funders":"Pontifícia Universidade Católica do Rio de Janeiro","keywords":"Computer science; Identification (biology)","score_opus":0.018437114190845247,"score_gpt":0.3041500359130949,"score_spread":0.28571292172224966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236803832","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95340353,0.000671615,0.04074345,0.0003091759,0.000042569398,0.00017346852,0.0003154278,0.0016395297,0.0027011805],"genre_scores_gemma":[0.9698074,0.00031773472,0.028709289,0.000034858065,0.0000093546705,0.000056436227,0.000379317,0.00014089311,0.0005446633],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9916808,0.0025530122,0.0012789414,0.0012754851,0.0026639672,0.0005477957],"domain_scores_gemma":[0.9485944,0.023664972,0.010083308,0.0053194724,0.011555158,0.0007825746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0085028205,0.00082366006,0.0009305563,0.0060195685,0.0010305587,0.0036277387,0.0012103791,0.0011582689,0.0013118414],"category_scores_gemma":[0.07022924,0.00054928963,0.00071327446,0.0037239373,0.0015638034,0.0028877296,0.0027449147,0.0011100121,0.00043285117],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052376575,0.0002658885,0.61245996,0.0007624971,0.00018070135,0.0032650947,0.021682054,0.003662084,0.022170844,0.0029982144,0.0015172635,0.3305117],"study_design_scores_gemma":[0.000064789936,0.0009357888,0.6651567,0.0010244104,0.00095756655,0.015623886,0.035303913,0.17151327,0.07843887,0.011822097,0.018879635,0.00027915],"about_ca_topic_score_codex":0.0061603035,"about_ca_topic_score_gemma":0.007200381,"teacher_disagreement_score":0.0085028205,"about_ca_system_score_codex":0.0013556264,"about_ca_system_score_gemma":0.0013799071,"threshold_uncertainty_score":0.04496777},"labels":[],"label_agreement":null},{"id":"W4237317117","doi":"10.1109/iwsc.2012.6227859","title":"Foreword","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Code refactoring; Software engineering; clone (Java method); Restructuring; Computer science; Software; Software development; Software system; Engineering management; Engineering; Programming language","score_opus":0.0237740953132858,"score_gpt":0.275310043769123,"score_spread":0.2515359484558372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4237317117","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00090021547,0.008790639,0.0043738373,0.06886587,0.13218199,0.00046076282,0.004200358,0.0019381236,0.77828825],"genre_scores_gemma":[0.0033934151,0.003281942,0.0013638064,0.01615477,0.010398874,0.0001743747,0.0024107595,0.00079724775,0.9620249],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99877423,0.00016588745,0.0000951228,0.0002316201,0.00059815677,0.0001349775],"domain_scores_gemma":[0.9967548,0.00047600825,0.00012898864,0.00028823793,0.0016404354,0.0007114085],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011936579,0.0010121665,0.0008991324,0.0015586658,0.0020287065,0.004910429,0.0018183871,0.0033579124,0.59790796],"category_scores_gemma":[0.007814994,0.00031758842,0.0007314556,0.001356175,0.0008283605,0.004879854,0.002903474,0.0039029578,0.55243766],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020730788,0.0000139559925,0.0000549389,0.00008835352,0.0000016902178,0.0000462433,0.000056886904,0.000015842226,0.000114069044,0.0045834207,0.9521434,0.042860497],"study_design_scores_gemma":[0.0000029435696,0.000008916531,0.000081242986,0.000067211324,0.0000010645368,0.000075031654,0.00004712102,0.000011331542,0.00004481887,0.0010188059,0.99863917,0.000002361974],"about_ca_topic_score_codex":0.0019692641,"about_ca_topic_score_gemma":0.0026973893,"teacher_disagreement_score":0.40209204,"about_ca_system_score_codex":0.0017865202,"about_ca_system_score_gemma":0.002295806,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4237331277","doi":"10.1145/944889.944891","title":"Designing UML diagrams for technical documentation","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Computer science; Unified Modeling Language; Documentation; UML tool; Applications of UML; Software engineering; Activity diagram; Class diagram; Use Case Diagram; Technical documentation; Communication diagram; Programming language; Software","score_opus":0.025381295572755046,"score_gpt":0.3086858361628847,"score_spread":0.28330454059012966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4237331277","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005024496,0.00052868837,0.9911732,0.0005673505,0.00019942626,0.00037304198,0.00018654135,0.0019657905,0.004503495],"genre_scores_gemma":[0.005499856,0.0006336901,0.9892212,0.00014621328,0.000078663106,0.000539679,0.0005529402,0.00054851605,0.0027792924],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.96407914,0.020308435,0.0044626202,0.0024963906,0.007703098,0.00095028983],"domain_scores_gemma":[0.96697783,0.015528131,0.0029074238,0.005006083,0.008830023,0.0007505037],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032731794,0.0028493109,0.0014486691,0.009859641,0.002981319,0.008907502,0.005279092,0.0042755767,0.010981462],"category_scores_gemma":[0.05011367,0.0027779844,0.0025059113,0.005747733,0.0030652448,0.0122375395,0.005597025,0.005385681,0.008108793],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007405946,0.00010128001,0.0008224679,0.0012468523,0.0000683178,0.0006947516,0.008922013,0.014438621,0.006449588,0.6096256,0.031626005,0.32593045],"study_design_scores_gemma":[0.000069729635,0.00007567492,0.00020614587,0.0011769043,0.00007028659,0.00069413224,0.0010596382,0.02199626,0.0054029357,0.14179073,0.82732916,0.00012841987],"about_ca_topic_score_codex":0.004938756,"about_ca_topic_score_gemma":0.0057507916,"teacher_disagreement_score":0.032731794,"about_ca_system_score_codex":0.0037037602,"about_ca_system_score_gemma":0.006505632,"threshold_uncertainty_score":0.17310429},"labels":[],"label_agreement":null},{"id":"W4237553142","doi":"10.4018/978-1-59140-941-1.ch002","title":"Intelligent Analysis of Software Maintenance Data","year":2007,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software maintenance; Software; Software development; Software engineering; Process (computing); Data mining; Data extraction; Software development process; Software sizing; Software construction","score_opus":0.05946106236650268,"score_gpt":0.3101298709114057,"score_spread":0.25066880854490303,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4237553142","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26559025,0.0014448696,0.7085358,0.00083257153,0.000087438755,0.00047878936,0.006726465,0.0103806155,0.005923248],"genre_scores_gemma":[0.52356005,0.0007390547,0.4601282,0.00012296873,0.00006031054,0.000283985,0.013002163,0.00022880417,0.0018745281],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99766195,0.00042998855,0.00025564214,0.00055203965,0.0009986309,0.00010169897],"domain_scores_gemma":[0.9943229,0.0027601658,0.0005934994,0.00081153674,0.0014262078,0.000085701016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002511732,0.00057061994,0.0009811013,0.0058769584,0.0004475349,0.0017594678,0.0007760149,0.00052879355,0.0009137218],"category_scores_gemma":[0.010343893,0.00031552804,0.0009010651,0.0029079376,0.00022454286,0.0014532738,0.0009896802,0.00081719976,0.0005676789],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030488102,0.00018916489,0.035897925,0.0004143638,0.0002011312,0.00037434115,0.000706955,0.059671994,0.017024197,0.004327313,0.005083139,0.8758045],"study_design_scores_gemma":[0.000017855935,0.00014437365,0.03761919,0.000109523025,0.000108654865,0.0002909952,0.0003431884,0.91818875,0.019561041,0.010416305,0.0131488005,0.00005139722],"about_ca_topic_score_codex":0.00229749,"about_ca_topic_score_gemma":0.0032119497,"teacher_disagreement_score":0.0058769584,"about_ca_system_score_codex":0.0007111863,"about_ca_system_score_gemma":0.0008393339,"threshold_uncertainty_score":0.0132834315},"labels":[],"label_agreement":null},{"id":"W4237642263","doi":"10.1109/iwsc.2013.6613048","title":"How much really changes? A case study of firefox version evolution using a clone detector","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"clone (Java method); Detector; Computer science; Code (set theory); Software evolution; Software; Programming language; Biology; Software development; Set (abstract data type); Gene; Telecommunications; Genetics","score_opus":0.03281116254218356,"score_gpt":0.2684347818150926,"score_spread":0.23562361927290906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4237642263","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99523395,0.00012155116,0.00380682,0.000108877764,0.000006023335,0.000041438405,0.000089146524,0.00006841303,0.0005238368],"genre_scores_gemma":[0.98159677,0.00010802383,0.017293088,0.00005159594,0.000008907225,0.000026969075,0.00023841027,0.000060590442,0.00061578985],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9960328,0.0016666105,0.00026112693,0.00065309607,0.0011371453,0.00024927626],"domain_scores_gemma":[0.9574653,0.033120103,0.0029266078,0.0029056515,0.002788065,0.0007942352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005712915,0.00037086126,0.00041231123,0.0023374143,0.0011795716,0.0010415482,0.0010049734,0.0011204899,0.0003338025],"category_scores_gemma":[0.022687932,0.00031113302,0.0005643457,0.0022749396,0.0010344845,0.0013717918,0.00076175726,0.0011357069,0.00008021847],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092489977,0.0015850062,0.6372669,0.00056901656,0.00038795118,0.019701501,0.03427148,0.036547832,0.057585917,0.005503268,0.0028774638,0.20277876],"study_design_scores_gemma":[0.00019315572,0.002313753,0.6868048,0.0001632838,0.00041122618,0.0118596,0.012003718,0.2033402,0.061567564,0.0045345863,0.016521094,0.0002870553],"about_ca_topic_score_codex":0.009706941,"about_ca_topic_score_gemma":0.014771474,"teacher_disagreement_score":0.009706941,"about_ca_system_score_codex":0.0010879413,"about_ca_system_score_gemma":0.00057972746,"threshold_uncertainty_score":0.030213118},"labels":[],"label_agreement":null},{"id":"W4237831168","doi":"10.1145/581388.581390","title":"Concern graphs","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Java; Source code; Scalability; Software engineering; Usability; Abstraction; Representation (politics); Graph; Software; Programming language; Theoretical computer science; Human–computer interaction; Database","score_opus":0.04323818430936585,"score_gpt":0.2629492601822788,"score_spread":0.21971107587291297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4237831168","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01160749,0.0008435344,0.9207001,0.0013111943,0.0002331919,0.0006774925,0.008826366,0.01106708,0.044733644],"genre_scores_gemma":[0.19353415,0.00237657,0.6875424,0.0012272702,0.0002076762,0.0013297091,0.036637407,0.0047750864,0.07236979],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99816316,0.0004144685,0.00018563322,0.00041648193,0.0006515324,0.00016864519],"domain_scores_gemma":[0.99630475,0.0013770957,0.00034038114,0.000850621,0.0009467578,0.000180404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012609937,0.0011715189,0.0004943518,0.0033039073,0.0015515995,0.002391771,0.001572347,0.0013946221,0.016551415],"category_scores_gemma":[0.006493564,0.00068978406,0.001667495,0.0026783617,0.00085680437,0.005417434,0.0023076378,0.0014802509,0.0046930593],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019937758,0.0001960422,0.0057292064,0.00099678,0.00015124088,0.0012835958,0.0019688352,0.014855822,0.0072098263,0.55165964,0.09424629,0.32150328],"study_design_scores_gemma":[0.0000282188,0.00006582665,0.0015654567,0.00017777213,0.000106568164,0.0010929556,0.00048067048,0.025100652,0.0060690283,0.2986592,0.66659796,0.000055728688],"about_ca_topic_score_codex":0.0072076274,"about_ca_topic_score_gemma":0.009048946,"teacher_disagreement_score":0.016551415,"about_ca_system_score_codex":0.0010952934,"about_ca_system_score_gemma":0.0014647124,"threshold_uncertainty_score":0.055369973},"labels":[],"label_agreement":null},{"id":"W4237865316","doi":"10.1109/icse.2015.238","title":"Mining Temporal Properties of Data Invariants","year":2015,"lang":"en","type":"article","venue":"2015 IEEE/ACM 37th IEEE International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Correctness; Temporal logic; Computer science; Property (philosophy); Linear temporal logic; Field (mathematics); Theoretical computer science; Value (mathematics); Data mining; Programming language; Mathematics; Machine learning","score_opus":0.25435898629449305,"score_gpt":0.35119945819126736,"score_spread":0.0968404718967743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4237865316","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27668622,0.000538461,0.7060137,0.0005689088,0.00005157621,0.00033660844,0.005269606,0.008044349,0.0024905403],"genre_scores_gemma":[0.7016611,0.00023943713,0.2884126,0.00014415792,0.0000323843,0.00030697603,0.0074964454,0.00057418534,0.0011327324],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9962547,0.00044363886,0.00040340037,0.0008381171,0.0017801286,0.00027995778],"domain_scores_gemma":[0.9763693,0.011265029,0.004498394,0.004037541,0.003477358,0.00035228537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002381469,0.0007963679,0.00058450253,0.0040940167,0.000628989,0.0013080431,0.0012222641,0.000611093,0.000894634],"category_scores_gemma":[0.023374896,0.0006932486,0.0014848963,0.0022828097,0.0012681156,0.0029957956,0.0013166963,0.0013798729,0.00034435614],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007297933,0.00065043505,0.21064676,0.0017819174,0.0003663896,0.0024591163,0.0023003977,0.14800985,0.06926886,0.06667488,0.008441475,0.48867005],"study_design_scores_gemma":[0.000060555092,0.00026478525,0.01580173,0.000130794,0.00014945837,0.0006464787,0.0005056066,0.833238,0.074725874,0.06374277,0.01065608,0.000077918245],"about_ca_topic_score_codex":0.0057317787,"about_ca_topic_score_gemma":0.009716328,"teacher_disagreement_score":0.0057317787,"about_ca_system_score_codex":0.0013639851,"about_ca_system_score_gemma":0.0030905865,"threshold_uncertainty_score":0.012594581},"labels":[],"label_agreement":null},{"id":"W4238014149","doi":"10.1002/(issn)1097-024x","title":"Software: Practice and Experience","year":2006,"lang":"en","type":"paratext","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":1222,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Software engineering; Computer science; Coding (social sciences); Software; Software development; World Wide Web; Data science; Operating system","score_opus":0.01775938229433564,"score_gpt":0.3186743183254481,"score_spread":0.30091493603111247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238014149","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0076748454,0.13384725,0.016411474,0.053793106,0.011981161,0.000083623636,0.00024829377,0.0012849203,0.7746753],"genre_scores_gemma":[0.04798666,0.074476615,0.0037458758,0.005221722,0.006726811,0.00014077664,0.00035337044,0.000927415,0.8604208],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970475,0.000959662,0.00022855746,0.0004523738,0.0010343678,0.00027750857],"domain_scores_gemma":[0.99444443,0.0016855774,0.00021689071,0.000512054,0.0015961741,0.0015448586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040483256,0.0018110013,0.0017967613,0.003886318,0.0033278326,0.017241081,0.0013046395,0.005066013,0.1405759],"category_scores_gemma":[0.008018508,0.0007542745,0.00057856966,0.0058601443,0.0047776136,0.015606838,0.008240465,0.004804713,0.075075716],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034376717,0.00009499649,0.00053493783,0.0009079705,0.000009736002,0.00023246536,0.0069256416,0.0002917197,0.0013395911,0.06535807,0.5847165,0.339554],"study_design_scores_gemma":[0.0000051285706,0.000030561965,0.0005031182,0.0006710716,0.0000034177108,0.00024291083,0.002029678,0.00012791775,0.00009376842,0.0075402227,0.9887416,0.000010558007],"about_ca_topic_score_codex":0.00078414194,"about_ca_topic_score_gemma":0.0018232368,"teacher_disagreement_score":0.1405759,"about_ca_system_score_codex":0.0017492283,"about_ca_system_score_gemma":0.002470998,"threshold_uncertainty_score":0.47027326},"labels":[],"label_agreement":null},{"id":"W4238050425","doi":"10.22215/etd/2014-10300","title":"A Comparative Study of Invariants Generated by Daikon and User-Defined Design Contracts","year":2014,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Complement (music); Quality (philosophy); Computer science; Chemistry","score_opus":0.04531319281334101,"score_gpt":0.3085843457145638,"score_spread":0.26327115290122277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238050425","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8771446,0.00043429193,0.11146406,0.00015463047,0.000031095733,0.00015227048,0.0006317813,0.0074914386,0.0024959866],"genre_scores_gemma":[0.8601865,0.00020528295,0.13506056,0.0000763497,0.000006722373,0.00010881949,0.0023307407,0.0010903905,0.0009345755],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98831546,0.0033322976,0.0009833544,0.0012120692,0.0055973264,0.00055943325],"domain_scores_gemma":[0.9108886,0.06422279,0.0069560497,0.011865517,0.0056738276,0.00039315587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008986393,0.0005297304,0.00035816542,0.0019300019,0.00033689226,0.0012541644,0.0013542365,0.000975475,0.0011795581],"category_scores_gemma":[0.051260866,0.00062427897,0.0007805649,0.001113636,0.0009911186,0.0017795678,0.0010927665,0.0010386958,0.0001978719],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026246998,0.0013363684,0.17085867,0.0022546197,0.00042348632,0.0017244914,0.00563859,0.22867097,0.12236849,0.01778358,0.0026705826,0.4436455],"study_design_scores_gemma":[0.0001597484,0.001431295,0.05280025,0.0002517052,0.00020154974,0.0010440595,0.0012775309,0.74582464,0.1797292,0.0051048803,0.012033543,0.00014158756],"about_ca_topic_score_codex":0.0027231858,"about_ca_topic_score_gemma":0.003918643,"teacher_disagreement_score":0.008986393,"about_ca_system_score_codex":0.0011305435,"about_ca_system_score_gemma":0.0013668793,"threshold_uncertainty_score":0.047525167},"labels":[],"label_agreement":null},{"id":"W4238083912","doi":"10.1109/msr.2015.12","title":"Co-evolution of Infrastructure and Source Code - An Empirical Study","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Server; Source code; Empirical research; Process (computing); Code (set theory); Database; Software engineering; Programming language; Operating system","score_opus":0.03351936749888526,"score_gpt":0.3357286193571814,"score_spread":0.30220925185829617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238083912","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9972832,0.00029572315,0.00063800946,0.00014657307,0.000005758545,0.00004699739,0.00014221302,0.000019368701,0.0014221302],"genre_scores_gemma":[0.9989976,0.00007944674,0.00040817773,0.000027478534,0.0000072889375,0.000034586315,0.00018977138,0.000017133294,0.00023850588],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9746636,0.013339989,0.0020532543,0.003376049,0.00511882,0.0014481914],"domain_scores_gemma":[0.4495099,0.42882878,0.07168853,0.02004262,0.021957204,0.007972917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01701433,0.000348054,0.0005195527,0.006336492,0.0013523168,0.0031946383,0.0019940957,0.001621851,0.003294844],"category_scores_gemma":[0.1571063,0.00062804733,0.0006390736,0.007194336,0.0033477352,0.0055334144,0.0035955464,0.003071586,0.00077914144],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000057966958,0.0001987184,0.993353,0.000024576586,0.000053830765,0.00012385321,0.0022004296,0.0001113292,0.00008511211,0.00017254996,0.00010126213,0.0035174068],"study_design_scores_gemma":[0.000008464688,0.00023849648,0.98873246,0.00004989708,0.000056007455,0.000659601,0.0065437714,0.0022342536,0.00023557228,0.00021758034,0.001005533,0.000018265462],"about_ca_topic_score_codex":0.0062057837,"about_ca_topic_score_gemma":0.006154861,"teacher_disagreement_score":0.01701433,"about_ca_system_score_codex":0.0015493039,"about_ca_system_score_gemma":0.0015186064,"threshold_uncertainty_score":0.08998144},"labels":[],"label_agreement":null},{"id":"W4238181319","doi":"10.1109/iwpse.2004.1334771","title":"Studying the evolution of software systems using evolutionary code extractors","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Computer science; Source code; Software engineering; Software system; Software; Code (set theory); KPI-driven code analysis; Software development; Software construction; Programming language","score_opus":0.03897183920478301,"score_gpt":0.27819398529890443,"score_spread":0.23922214609412143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238181319","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38149664,0.00041844923,0.61274755,0.00039349124,0.000020718058,0.0002284749,0.000725162,0.0019846684,0.0019848456],"genre_scores_gemma":[0.41197592,0.0005755503,0.58386934,0.00005515367,0.00001701988,0.00017720785,0.0018525787,0.00032287725,0.0011543662],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984534,0.0005670979,0.00014302904,0.00023023158,0.00054537586,0.000060812097],"domain_scores_gemma":[0.98124146,0.012769923,0.0021290726,0.0019735927,0.0017037275,0.00018220137],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021059345,0.0004963453,0.0003949623,0.0038735396,0.0005945238,0.000979766,0.0007181886,0.00080880197,0.0006564285],"category_scores_gemma":[0.020869736,0.00042218278,0.00063158525,0.0034119163,0.00063521514,0.0022244428,0.0006488112,0.001103762,0.00024162109],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002298271,0.00041569193,0.10408193,0.0006802464,0.00022859468,0.00094185985,0.003290789,0.10156592,0.07233946,0.018562274,0.0013911028,0.69627225],"study_design_scores_gemma":[0.00007421584,0.00038212206,0.072446115,0.00017751337,0.00018924297,0.0011095834,0.0008688024,0.7583062,0.12611549,0.027647454,0.012559472,0.00012381634],"about_ca_topic_score_codex":0.0022429468,"about_ca_topic_score_gemma":0.0033877075,"teacher_disagreement_score":0.0038735396,"about_ca_system_score_codex":0.00045860352,"about_ca_system_score_gemma":0.0006698533,"threshold_uncertainty_score":0.011137366},"labels":[],"label_agreement":null},{"id":"W4238222639","doi":"10.1145/354222.353194","title":"An Aristotelian understanding of object-oriented programming","year":2000,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Jerome's University; University of Waterloo","funders":"","keywords":"Computer science; Object (grammar); Predicate (mathematical logic); Programming language; Logic programming; Object-oriented programming; Programming paradigm; First-order logic; Relation (database); Inductive programming; Epistemology; Artificial intelligence; Philosophy","score_opus":0.03571274590944432,"score_gpt":0.28869754795557623,"score_spread":0.2529848020461319,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238222639","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014743818,0.019712523,0.42550772,0.041239683,0.001260666,0.000063387844,0.00014205213,0.00038823357,0.4969419],"genre_scores_gemma":[0.659008,0.018903555,0.22515336,0.013973322,0.004037504,0.00039041706,0.0002828076,0.00033875214,0.077912256],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9972801,0.00083775487,0.00011465908,0.00040738677,0.0011172142,0.00024293683],"domain_scores_gemma":[0.9983901,0.0007212104,0.00012856824,0.00026100018,0.00034790364,0.00015125208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032592476,0.0006287897,0.0007918499,0.0022972538,0.004293629,0.0086639,0.0018659257,0.0037279213,0.0034537006],"category_scores_gemma":[0.003654547,0.00046541024,0.00080083107,0.0019339011,0.021455254,0.01399536,0.0025859333,0.005139453,0.0014269917],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000017787869,0.0000037281009,0.000018900495,0.000007846079,6.5106633e-7,0.000010850826,0.00031481305,0.00009034102,0.000028522323,0.9973707,0.000629805,0.0015220789],"study_design_scores_gemma":[0.0000031219683,0.0000042733336,0.00004000886,0.000020741028,0.0000014139788,0.00002315977,0.00008538513,0.0003097369,0.000036205503,0.97127354,0.028198896,0.000003397521],"about_ca_topic_score_codex":0.0031747662,"about_ca_topic_score_gemma":0.0024045156,"teacher_disagreement_score":0.0086639,"about_ca_system_score_codex":0.0039572,"about_ca_system_score_gemma":0.0027352455,"threshold_uncertainty_score":0.028711617},"labels":[],"label_agreement":null},{"id":"W4238417811","doi":"10.1109/icse.2003.1201304","title":"FEAT a tool for locating, describing, and analyzing concerns in source code","year":2003,"lang":"en","type":"article","venue":"25th International Conference on Software Engineering, 2003. Proceedings.","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Source code; Task (project management); Code (set theory); Programming language; Software engineering; Code review; KPI-driven code analysis; Open source; Static program analysis; Software development; Software; Systems engineering; Engineering","score_opus":0.05235386629265224,"score_gpt":0.29129267946729115,"score_spread":0.23893881317463891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238417811","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0048420103,0.00048672027,0.7262996,0.00024077056,0.00012984363,0.0004000089,0.0044078073,0.25755563,0.0056375773],"genre_scores_gemma":[0.045150016,0.0008788231,0.87848055,0.00037793114,0.000092378104,0.0009864848,0.02236527,0.03605144,0.015617109],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99812824,0.00039714287,0.00024705814,0.0003069327,0.00077713944,0.00014354552],"domain_scores_gemma":[0.99242485,0.004633057,0.0007637523,0.0012549461,0.00078235014,0.0001409545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029898083,0.0032065227,0.0010491137,0.007514638,0.0019055858,0.003243704,0.0029738704,0.0028549223,0.020302106],"category_scores_gemma":[0.014221382,0.0026804516,0.00229287,0.0025863785,0.0014932425,0.0070186383,0.0036722017,0.003253773,0.010013874],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007338863,0.0005725955,0.015934482,0.0033665437,0.00051685405,0.0024131255,0.0044277245,0.0135106575,0.033591628,0.057239402,0.33583528,0.53185785],"study_design_scores_gemma":[0.00042232074,0.00046260466,0.008530466,0.0012371505,0.00036743196,0.0050790496,0.0009914377,0.15852979,0.05467959,0.057807725,0.71140265,0.0004897621],"about_ca_topic_score_codex":0.0058216075,"about_ca_topic_score_gemma":0.008650236,"teacher_disagreement_score":0.020302106,"about_ca_system_score_codex":0.0011694607,"about_ca_system_score_gemma":0.0023259725,"threshold_uncertainty_score":0.06791735},"labels":[],"label_agreement":null},{"id":"W4238802097","doi":"10.22215/etd/2018-13442","title":"Traceability Modeling for the Engineering of Heterogeneous Systems","year":2018,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Consortium de Recherche et d’innovation en Aérospatiale au Québec; University of Ottawa","keywords":"Traceability; Requirements traceability; Computer science; Unified Modeling Language; Software engineering; Model-driven architecture; Context (archaeology); Systems engineering; Requirements engineering; Programming language; Engineering; Software; Requirement","score_opus":0.026323723713719206,"score_gpt":0.2812324750236669,"score_spread":0.25490875130994767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238802097","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017398723,0.0012496803,0.99077016,0.00048717757,0.000059122543,0.00010709162,0.000108626395,0.00039813502,0.005080125],"genre_scores_gemma":[0.1328043,0.008504561,0.84325755,0.0003766173,0.00025389311,0.0009344755,0.0013858174,0.00067509274,0.011807707],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965373,0.0012398139,0.00034711452,0.00046326756,0.0012539452,0.00015860658],"domain_scores_gemma":[0.995974,0.002323832,0.00043455823,0.0007254337,0.00045775966,0.00008444549],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035733604,0.0018737921,0.0007212706,0.0033643031,0.0010328104,0.0031991166,0.002310427,0.0021811684,0.00439672],"category_scores_gemma":[0.009243824,0.00083239115,0.0027818296,0.003379722,0.0017817299,0.005550854,0.0024225495,0.003426525,0.0012884174],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000036166985,0.00012778935,0.0011994052,0.00069364754,0.00014712996,0.00031235462,0.0009275522,0.18636194,0.003880268,0.714269,0.0030158532,0.089028895],"study_design_scores_gemma":[0.000024530866,0.000071942115,0.00042915132,0.00064385834,0.000118991135,0.00026032943,0.00020879142,0.4021967,0.00430639,0.51598233,0.07570153,0.000055535325],"about_ca_topic_score_codex":0.009384567,"about_ca_topic_score_gemma":0.0066018105,"teacher_disagreement_score":0.009384567,"about_ca_system_score_codex":0.0030240314,"about_ca_system_score_gemma":0.0035671892,"threshold_uncertainty_score":0.021940947},"labels":[],"label_agreement":null},{"id":"W4238850792","doi":"10.1145/1449955.1449786","title":"The impact of static-dynamic coupling on remodularization","year":2008,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Static analysis; Modular design; Coupling (piping); Java; Programming language; Process (computing); Point (geometry); Mathematics","score_opus":0.02526368840728023,"score_gpt":0.30369978492059846,"score_spread":0.27843609651331824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238850792","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9615411,0.00023209685,0.033735413,0.00020003383,0.000012840697,0.00008604118,0.000035934885,0.00028953704,0.0038669086],"genre_scores_gemma":[0.9920169,0.000051201332,0.0075095934,0.000033928653,0.000006170616,0.0000258073,0.000032719196,0.00008606946,0.00023759055],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9834719,0.007023061,0.0010618374,0.0015278048,0.0053586005,0.0015568114],"domain_scores_gemma":[0.73039687,0.21202725,0.022279454,0.020701936,0.011473732,0.0031206599],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014978382,0.0007340162,0.0003700222,0.0018886934,0.000976958,0.002231232,0.0011529289,0.00072717154,0.0021912593],"category_scores_gemma":[0.123060055,0.00053814845,0.0005625391,0.0010271017,0.0024202908,0.0043740286,0.0028476152,0.001382538,0.00019318535],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00164189,0.0013098087,0.4003324,0.0014139863,0.00049414573,0.0021065413,0.011058997,0.1328238,0.11751828,0.030554162,0.0008023064,0.29994377],"study_design_scores_gemma":[0.00020825886,0.0066320724,0.4499261,0.00048989797,0.0010242359,0.0038448905,0.010694761,0.2990703,0.17178254,0.044872195,0.01102094,0.00043379466],"about_ca_topic_score_codex":0.0017347133,"about_ca_topic_score_gemma":0.0021979222,"teacher_disagreement_score":0.014978382,"about_ca_system_score_codex":0.0010975328,"about_ca_system_score_gemma":0.0012555096,"threshold_uncertainty_score":0.079214215},"labels":[],"label_agreement":null},{"id":"W4239042842","doi":"10.22215/etd/2011-09062","title":"Reverse Engineering of Java programs through static and dynamic analysis to generate scenario diagrams","year":2011,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Heritage; Library and Archives Canada","funders":"","keywords":"Computer science; Java; Programming language","score_opus":0.01643957846396706,"score_gpt":0.2670372798167705,"score_spread":0.25059770135280346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239042842","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024611427,0.00011789518,0.9630268,0.0002016098,0.000038190883,0.00030550506,0.0003432065,0.00616089,0.005194317],"genre_scores_gemma":[0.14636096,0.00024797712,0.8461599,0.00008471964,0.000010810638,0.00028084512,0.0018539457,0.0013486854,0.0036520893],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99818987,0.00062335,0.000106583844,0.00020749333,0.00074233295,0.00013042359],"domain_scores_gemma":[0.99354035,0.0041156104,0.0003498419,0.00086184667,0.0010491648,0.00008322598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017859567,0.00076195184,0.0003900792,0.0025352766,0.0005545291,0.0013703371,0.00093576463,0.00075474294,0.0040979185],"category_scores_gemma":[0.009274762,0.00062359276,0.0013891787,0.0009820696,0.00066390954,0.0013723877,0.0012191631,0.0012172062,0.0014757555],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020307917,0.0006101364,0.008182714,0.00088637276,0.00016353061,0.0012571913,0.0016846174,0.18768105,0.058175948,0.09588355,0.009457123,0.6358147],"study_design_scores_gemma":[0.00007094001,0.00015087883,0.0015540711,0.0001910932,0.00012758454,0.00062029867,0.0003559145,0.8295248,0.079041824,0.042078838,0.04622029,0.000063449435],"about_ca_topic_score_codex":0.0044025974,"about_ca_topic_score_gemma":0.0092594875,"teacher_disagreement_score":0.0044025974,"about_ca_system_score_codex":0.00087337184,"about_ca_system_score_gemma":0.002244753,"threshold_uncertainty_score":0.0137088895},"labels":[],"label_agreement":null},{"id":"W4239191916","doi":"10.1145/2398857.2384665","title":"Speculative analysis of integrated development environment recommendations","year":2012,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Code refactoring; Source code; Eclipse; Programming language; Software engineering; Code (set theory); Domain (mathematical analysis); Semantics (computer science); Executable; Software","score_opus":0.040113989935213284,"score_gpt":0.2883590507156111,"score_spread":0.2482450607803978,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239191916","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.59514713,0.0006156454,0.38627312,0.0014215839,0.00010713882,0.0006049414,0.0021741514,0.004565384,0.009090808],"genre_scores_gemma":[0.8659952,0.00015354104,0.12927623,0.0001384952,0.00003951007,0.00017867963,0.0020526752,0.00034524835,0.0018203834],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.992652,0.0022996608,0.00036513474,0.0008844111,0.0033245913,0.0004741547],"domain_scores_gemma":[0.9202562,0.06312679,0.0029101712,0.0058225766,0.0073274556,0.00055677834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059188385,0.0008404732,0.00062872743,0.0020382844,0.00060441985,0.0018881471,0.0013652191,0.0008413342,0.0036006165],"category_scores_gemma":[0.086629204,0.0008190723,0.00079799467,0.0018155565,0.0006293437,0.0027151145,0.0012054611,0.0014609192,0.0004948777],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002558933,0.00050471607,0.19166587,0.00087638607,0.0004931901,0.0013737194,0.0021737607,0.40349036,0.019894809,0.034912102,0.013316129,0.32874006],"study_design_scores_gemma":[0.000065686465,0.000099962264,0.011673769,0.000032581065,0.000067492874,0.00008427595,0.00020279526,0.9727692,0.0036207854,0.00892878,0.002430666,0.000024148641],"about_ca_topic_score_codex":0.012323993,"about_ca_topic_score_gemma":0.017911397,"teacher_disagreement_score":0.012323993,"about_ca_system_score_codex":0.0013660294,"about_ca_system_score_gemma":0.001821659,"threshold_uncertainty_score":0.031302214},"labels":[],"label_agreement":null},{"id":"W4239365492","doi":"10.1109/iwpse.2004.1334766","title":"An automatic approach to identify class evolution discontinuities","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code refactoring; Software evolution; Computer science; TRACE (psycholinguistics); Java; Class (philosophy); Software; Software maintenance; Software system; Helpfulness; Software walkthrough; Programming language; Software engineering; Artificial intelligence; Software construction","score_opus":0.018720802162028528,"score_gpt":0.2935178736171462,"score_spread":0.27479707145511767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239365492","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09538275,0.0005044073,0.8888608,0.00020537095,0.00005123239,0.00039853886,0.0010356054,0.011385399,0.0021758436],"genre_scores_gemma":[0.3074109,0.00016773267,0.6882023,0.00006204594,0.000034537323,0.00031391473,0.0016781773,0.00032028937,0.0018100968],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975872,0.00031105327,0.00023674428,0.00066815235,0.0010297942,0.00016702153],"domain_scores_gemma":[0.99295276,0.0026553536,0.0013007263,0.00091370154,0.002043428,0.0001339918],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015149862,0.0006896263,0.0009978409,0.011812513,0.00081795384,0.0016481847,0.0017159047,0.0011953544,0.0022914251],"category_scores_gemma":[0.0081816735,0.00042029866,0.00073897804,0.0046786126,0.00041566332,0.0019599118,0.0011398043,0.0008281525,0.0011255061],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003529675,0.00024536203,0.014817232,0.00034179477,0.000106655585,0.0003792862,0.00077584194,0.0046431893,0.0759771,0.003138285,0.0036340896,0.8955883],"study_design_scores_gemma":[0.00026127417,0.0005639035,0.07436942,0.00016325827,0.00044994248,0.0027086567,0.001099028,0.7367062,0.14403477,0.013687941,0.02564331,0.0003123394],"about_ca_topic_score_codex":0.0041459524,"about_ca_topic_score_gemma":0.0042854724,"teacher_disagreement_score":0.011812513,"about_ca_system_score_codex":0.0007289939,"about_ca_system_score_gemma":0.0013138046,"threshold_uncertainty_score":0.00824362},"labels":[],"label_agreement":null},{"id":"W4239912276","doi":"10.1002/0471028959.sof198","title":"Measurement","year":2002,"lang":"en","type":"other","venue":"Encyclopedia of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"PricewaterhouseCoopers (Canada)","funders":"","keywords":"Computer science; Software engineering; Software; Software metric; Context (archaeology); Compiler; Software development; Process (computing); Resource (disambiguation); Software sizing; Software construction; Software system; Software measurement; Programming language","score_opus":0.013942480975880368,"score_gpt":0.21433137176665734,"score_spread":0.20038889079077699,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239912276","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050895475,0.0061752694,0.04965243,0.01802106,0.0056562442,0.0009407298,0.011496318,0.0023193818,0.900649],"genre_scores_gemma":[0.18025807,0.017351275,0.057322454,0.012387943,0.004119228,0.0024014371,0.028380152,0.0024769888,0.69530255],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98759395,0.0030506735,0.0009202428,0.0020244454,0.0057098567,0.0007008484],"domain_scores_gemma":[0.9837908,0.0025525326,0.00087250373,0.0032001783,0.008737528,0.00084644713],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008605278,0.001324065,0.001045001,0.0049939184,0.00228818,0.011278866,0.0028606073,0.0019911067,0.15229549],"category_scores_gemma":[0.032719668,0.00048792662,0.0010256051,0.006599806,0.0019006714,0.009839877,0.006006106,0.0027234447,0.0901324],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011021805,0.00010457481,0.005142684,0.000746424,0.00005237701,0.00010361733,0.0014653677,0.00062871695,0.00044768958,0.2962136,0.29122844,0.40375635],"study_design_scores_gemma":[0.000011335581,0.00004768281,0.0032203265,0.00059799326,0.000018271676,0.00012738028,0.0009469941,0.00037710514,0.00038522948,0.03315101,0.96108705,0.000029562298],"about_ca_topic_score_codex":0.0049788007,"about_ca_topic_score_gemma":0.0024812345,"teacher_disagreement_score":0.15229549,"about_ca_system_score_codex":0.005209568,"about_ca_system_score_gemma":0.006490393,"threshold_uncertainty_score":0.50947917},"labels":[],"label_agreement":null},{"id":"W4239921431","doi":"10.1109/icse.2003.1201214","title":"Design pattern rationale graphs: linking design to source","year":2003,"lang":"en","type":"article","venue":"25th International Conference on Software Engineering, 2003. Proceedings.","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Software design pattern; Computer science; Design pattern; Structural pattern; Flexibility (engineering); Pattern language (formal languages); Source code; TRACE (psycholinguistics); Code (set theory); Representation (politics); Software engineering; Human–computer interaction; Programming language; Software design; Software development; Software","score_opus":0.05281837649059293,"score_gpt":0.2709705174956941,"score_spread":0.2181521410051012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239921431","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021429071,0.00018363154,0.98307216,0.00059670646,0.00006740093,0.0004029993,0.00081523857,0.008832308,0.0038867404],"genre_scores_gemma":[0.022885852,0.00068951846,0.9661985,0.00036891413,0.000047226273,0.00076161354,0.0033918207,0.0025211354,0.0031354977],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934676,0.0028303522,0.0006144592,0.00067917403,0.0021989788,0.00020943438],"domain_scores_gemma":[0.9730483,0.016089791,0.0023185192,0.005203233,0.0029279913,0.00041208093],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007895775,0.0022589648,0.00073829835,0.00860691,0.001504263,0.004369649,0.002866488,0.0027972176,0.009306025],"category_scores_gemma":[0.044769086,0.002085609,0.001869139,0.0059175147,0.0021146857,0.00749288,0.004954023,0.003089373,0.003393876],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001856561,0.00039850784,0.0043287515,0.0019330947,0.00019954344,0.0014165007,0.0049978197,0.02851656,0.008240688,0.23138304,0.04905101,0.6693489],"study_design_scores_gemma":[0.00017285226,0.00015120344,0.0017980759,0.0012424266,0.00017681284,0.0011004183,0.0010857906,0.15819988,0.017102439,0.41695517,0.40178314,0.00023173462],"about_ca_topic_score_codex":0.007418426,"about_ca_topic_score_gemma":0.009502147,"teacher_disagreement_score":0.009306025,"about_ca_system_score_codex":0.0015291978,"about_ca_system_score_gemma":0.0040186155,"threshold_uncertainty_score":0.041757405},"labels":[],"label_agreement":null},{"id":"W4240056776","doi":"10.1002/spip.360","title":"A method for re‐planning of software releases using discrete‐event simulation","year":2008,"lang":"en","type":"article","venue":"Software Process Improvement and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Universität Ulm","keywords":"Planner; Computer science; Task (project management); Process (computing); Operational planning; Software release life cycle; Software; Operations research; Event (particle physics); Plan (archaeology); Phase (matter); Discrete event simulation; Work (physics); Product (mathematics); Realization (probability); Industrial engineering; Simulation; Systems engineering; Engineering; Software system; Artificial intelligence; Mathematics; Business","score_opus":0.07030822308933612,"score_gpt":0.41155360592307805,"score_spread":0.34124538283374195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4240056776","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032239573,0.000035663135,0.99504673,0.000049969178,0.000012723847,0.00006217573,0.00003474841,0.00086647837,0.0006675402],"genre_scores_gemma":[0.20169504,0.000085223655,0.7965379,0.00002992759,0.000014631129,0.00032383678,0.00019975366,0.0001266323,0.000986993],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985555,0.0005889732,0.000097663935,0.00025268446,0.00042906246,0.000076279546],"domain_scores_gemma":[0.9956285,0.0030953758,0.00030349125,0.0003947216,0.0004356064,0.0001423906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028982756,0.0011225743,0.0010683123,0.0013174578,0.00046113922,0.001519567,0.0018353655,0.0009719581,0.0033675872],"category_scores_gemma":[0.0066496427,0.000911197,0.0011695398,0.00078168296,0.00069911825,0.0011086965,0.0011368504,0.0014441519,0.00055737275],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008070322,0.000041214298,0.00054831477,0.000064121574,0.000051582225,0.000079916885,0.00009228559,0.94628656,0.0010585475,0.008052156,0.0004460038,0.043198626],"study_design_scores_gemma":[0.00001260764,0.000007851687,0.000023459548,0.0000048710126,0.000005639373,0.00000805495,0.000004367118,0.9976792,0.0003738598,0.0013188021,0.00055703276,0.000004439906],"about_ca_topic_score_codex":0.008790087,"about_ca_topic_score_gemma":0.0064375927,"teacher_disagreement_score":0.008790087,"about_ca_system_score_codex":0.0009997049,"about_ca_system_score_gemma":0.0015783514,"threshold_uncertainty_score":0.01747781},"labels":[],"label_agreement":null},{"id":"W4240298746","doi":"10.1145/1095430.1081736","title":"Facilitating software evolution research with kenyon","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software evolution; Source code; Software engineering; Software; Commit; Metadata; Resource (disambiguation); Software system; Software development; Set (abstract data type); Software construction; Data science; Database; Programming language; Operating system","score_opus":0.03927975784835953,"score_gpt":0.2942376257946139,"score_spread":0.25495786794625436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4240298746","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10621343,0.00042181805,0.67963463,0.0006349016,0.000083968254,0.000672082,0.0011950815,0.1968528,0.014291216],"genre_scores_gemma":[0.14723536,0.00032096147,0.83498806,0.00025356823,0.000021699296,0.0005471693,0.0014927625,0.008243075,0.006897348],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99784195,0.0005351109,0.00020632723,0.0005636159,0.0006751419,0.00017778447],"domain_scores_gemma":[0.9850483,0.0077743162,0.00078443676,0.0049084327,0.0010342787,0.00045031548],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039040202,0.0009420432,0.0009582919,0.002201976,0.0007894645,0.0020879027,0.0021792392,0.00073860027,0.010160705],"category_scores_gemma":[0.015691817,0.0017533796,0.000691908,0.0026562775,0.00092669996,0.005959814,0.006547942,0.0016244913,0.0031895195],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002458774,0.0006046484,0.03153631,0.0011719783,0.00037963063,0.0018183462,0.014173456,0.012008298,0.07657408,0.020190658,0.065561004,0.7735228],"study_design_scores_gemma":[0.0012810384,0.0012859724,0.036929414,0.000663281,0.0006847713,0.004223532,0.005377839,0.22252975,0.15919356,0.033557735,0.5335139,0.0007591505],"about_ca_topic_score_codex":0.0027502521,"about_ca_topic_score_gemma":0.010913928,"teacher_disagreement_score":0.010160705,"about_ca_system_score_codex":0.0005813243,"about_ca_system_score_gemma":0.0019341718,"threshold_uncertainty_score":0.03399098},"labels":[],"label_agreement":null},{"id":"W4240736797","doi":"10.1109/icse.2005.1553554","title":"Using structural context to recommend source code examples","year":2005,"lang":"en","type":"article","venue":"Proceedings. 27th International Conference on Software Engineering, 2005. ICSE 2005.","topic":"Software Engineering Research","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Programming language; Source code; Eclipse; Coding (social sciences); Code (set theory); Context (archaeology); Code review; Class (philosophy); Task (project management); Software engineering; Static program analysis; Artificial intelligence; Software development; Software","score_opus":0.06840004737661666,"score_gpt":0.31654824771238393,"score_spread":0.2481482003357673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4240736797","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53147113,0.0038574461,0.43431973,0.0013366725,0.00020300364,0.0009903619,0.004720969,0.011969677,0.0111309495],"genre_scores_gemma":[0.6418262,0.0009470137,0.34472927,0.00019073255,0.00008863647,0.00046111151,0.008907407,0.0005992816,0.0022503934],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.998519,0.00033483547,0.00012054525,0.00038761736,0.0005604683,0.00007740232],"domain_scores_gemma":[0.9872869,0.007771716,0.0008959124,0.0008917568,0.00282727,0.0003264352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013668635,0.0009236082,0.00074928324,0.0116271665,0.00087604386,0.00139657,0.0012108161,0.0016063306,0.0025870833],"category_scores_gemma":[0.020001138,0.0006529509,0.00060267514,0.004178875,0.00042225886,0.0024113872,0.0010401069,0.0010226619,0.0013198063],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073263975,0.0007410059,0.12460185,0.0012764687,0.00026590537,0.0011610078,0.0024397953,0.02147619,0.016258748,0.0043565105,0.029398013,0.79729193],"study_design_scores_gemma":[0.00041258454,0.0007061492,0.066861145,0.0008933516,0.00072440127,0.0022913879,0.002704463,0.807664,0.029998349,0.023345968,0.06413507,0.0002631363],"about_ca_topic_score_codex":0.009186104,"about_ca_topic_score_gemma":0.03188726,"teacher_disagreement_score":0.0116271665,"about_ca_system_score_codex":0.00062802935,"about_ca_system_score_gemma":0.0015768396,"threshold_uncertainty_score":0.018265247},"labels":[],"label_agreement":null},{"id":"W4240790600","doi":"10.1145/1932682.1869518","title":"Refactoring references for library migration","year":2010,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Code refactoring; Computer science; Programming language; Declaration; Set (abstract data type); Transformation (genetics); Source code; Feature (linguistics); Software engineering; Code (set theory); Key (lock); Program transformation; Software; Operating system","score_opus":0.03467563794292089,"score_gpt":0.28359976719964003,"score_spread":0.24892412925671914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4240790600","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06538603,0.0015716245,0.8656142,0.0010555681,0.00029973534,0.00037291335,0.00033131876,0.05839985,0.00696877],"genre_scores_gemma":[0.20724134,0.0007905479,0.7782799,0.00045900443,0.00010897238,0.0002910387,0.0009727634,0.0064272135,0.0054292437],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99253744,0.0027521763,0.00063735706,0.0008844085,0.0027493797,0.00043918184],"domain_scores_gemma":[0.95126504,0.017086627,0.004108234,0.021080457,0.005894777,0.0005648571],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071594156,0.0011654638,0.00080239243,0.0022937332,0.0014939422,0.0023196982,0.0038007386,0.0027405522,0.005759165],"category_scores_gemma":[0.046775308,0.0011440978,0.001278156,0.0021312789,0.0012697636,0.0058333594,0.004998804,0.002850133,0.0031746875],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042182003,0.0004679915,0.014351348,0.0013113735,0.00012546414,0.0013661913,0.0035018784,0.01840528,0.06596809,0.030353446,0.025325894,0.8384012],"study_design_scores_gemma":[0.00037077584,0.0007210265,0.010143742,0.0012880855,0.00047986035,0.0034488845,0.0011770029,0.22441438,0.25601393,0.056234013,0.44518045,0.0005279322],"about_ca_topic_score_codex":0.0023073787,"about_ca_topic_score_gemma":0.0038980632,"teacher_disagreement_score":0.0071594156,"about_ca_system_score_codex":0.0008876819,"about_ca_system_score_gemma":0.0031953508,"threshold_uncertainty_score":0.037863016},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"not_applicable","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"design_other","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W4241052220","doi":"10.1007/978-981-16-1927-4_7","title":"BigCloneBench","year":2021,"lang":"de","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Benchmark (surveying); clone (Java method); Benchmarking; Computer science; Variety (cybernetics); Process (computing); Data mining; Artificial intelligence; Programming language; Biology; Geography","score_opus":0.032350085397654126,"score_gpt":0.2589332223423216,"score_spread":0.22658313694466745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241052220","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017747746,0.0022663593,0.114146434,0.0011495365,0.001031248,0.00023440113,0.015621598,0.07724802,0.7865275],"genre_scores_gemma":[0.005927733,0.0016935803,0.051786043,0.0011811261,0.00024181195,0.0003086441,0.036786877,0.026799317,0.87527484],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99954563,0.00004951744,0.000020365896,0.00009218542,0.00025205853,0.000040236275],"domain_scores_gemma":[0.99923325,0.00023024667,0.000023165418,0.00021878563,0.00018736406,0.00010720082],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00053020974,0.0016770533,0.0008493039,0.0023666283,0.001041566,0.0037932554,0.0026192975,0.0010984876,0.42305198],"category_scores_gemma":[0.0020359948,0.0011279919,0.0009929906,0.0034548745,0.00040295685,0.0045038797,0.0024834508,0.0023016448,0.33131415],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037235644,0.00003825289,0.000092230315,0.00015789593,0.0000063696843,0.000041637086,0.0000676628,0.00029316437,0.0020244496,0.014735652,0.80687064,0.17563489],"study_design_scores_gemma":[0.000013783392,0.000017065287,0.00014798855,0.000071897244,0.0000078138255,0.00016055047,0.000030405925,0.0010799222,0.0023125077,0.012413267,0.9837298,0.000014981092],"about_ca_topic_score_codex":0.0013287815,"about_ca_topic_score_gemma":0.0034917966,"teacher_disagreement_score":0.42305198,"about_ca_system_score_codex":0.0008941219,"about_ca_system_score_gemma":0.00097420765,"threshold_uncertainty_score":0.8229463},"labels":[],"label_agreement":null},{"id":"W4241100723","doi":"10.1109/icse.2013.6606723","title":"Normalizing source code vocabulary to support program comprehension and software quality","year":2013,"lang":"en","type":"article","venue":"2013 35th International Conference on Software Engineering (ICSE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Program comprehension; Computer science; Source code; Identifier; Normalization (sociology); Vocabulary; Software maintenance; Natural language processing; Static program analysis; Software quality; Information retrieval; Software; Artificial intelligence; Programming language; Software development; Software system; Linguistics","score_opus":0.053262416941427956,"score_gpt":0.32160251450881083,"score_spread":0.26834009756738286,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241100723","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17160854,0.0016194214,0.7898248,0.0004970203,0.00010567147,0.0006731546,0.0007354515,0.03060301,0.0043329927],"genre_scores_gemma":[0.47098845,0.0006303929,0.51895607,0.00026255965,0.000073184034,0.0006028397,0.0029962359,0.0038780977,0.0016122382],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946331,0.0012504436,0.0007809167,0.0013131127,0.0018047253,0.00021775308],"domain_scores_gemma":[0.97338945,0.010545646,0.004239269,0.005797343,0.0056835776,0.0003446508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034929598,0.0010203599,0.00095528946,0.0031927135,0.00065039954,0.0021419637,0.0015893014,0.00073729607,0.0016442913],"category_scores_gemma":[0.044149783,0.0005078607,0.0008040629,0.0025966158,0.0010247105,0.005288298,0.0027561188,0.0014158633,0.0009929513],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041620972,0.0003891279,0.013429092,0.0013001462,0.000111106536,0.00021860999,0.0037694299,0.01122505,0.1751803,0.009189625,0.0051296074,0.7796417],"study_design_scores_gemma":[0.00022247345,0.00085224456,0.03625889,0.00063909934,0.00060102384,0.0014807734,0.0027447282,0.30303317,0.53481036,0.040307682,0.078687534,0.00036203762],"about_ca_topic_score_codex":0.0029236833,"about_ca_topic_score_gemma":0.0029597753,"teacher_disagreement_score":0.0034929598,"about_ca_system_score_codex":0.0011352526,"about_ca_system_score_gemma":0.0024090286,"threshold_uncertainty_score":0.01847279},"labels":[],"label_agreement":null},{"id":"W4241282950","doi":"10.1109/icpc.2015.11","title":"Detection of Software Evolution Phases Based on Development Activities","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Software evolution; Computer science; Software development; Granularity; Software; Commit; Software construction; Software engineering; Software sizing; Software analytics; Software maintenance; Data mining; Database; Programming language","score_opus":0.031700567589858446,"score_gpt":0.26260489484339994,"score_spread":0.23090432725354149,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241282950","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6682402,0.00088246394,0.32230896,0.00020250914,0.000038730093,0.0005101989,0.0017865107,0.0028283854,0.003202086],"genre_scores_gemma":[0.7384203,0.00035588356,0.25661975,0.000034497774,0.000022329948,0.00021368139,0.0030957686,0.0002253542,0.0010124724],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9978667,0.0003971488,0.0002601228,0.0004747443,0.0008571495,0.00014408852],"domain_scores_gemma":[0.9803094,0.00998837,0.0038449662,0.0016534857,0.0035936295,0.00061003433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020144954,0.00060965377,0.00054216944,0.008790963,0.00048620836,0.0012891858,0.0006334028,0.0005051406,0.0008227097],"category_scores_gemma":[0.016659766,0.00044638367,0.00062487065,0.003921864,0.00037599684,0.0011966075,0.00095403066,0.0007923594,0.00028295466],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006778195,0.00034357153,0.29169726,0.0008890392,0.00017610172,0.0007146315,0.0027568613,0.03042163,0.053171016,0.005901252,0.0026595425,0.61059123],"study_design_scores_gemma":[0.000114232935,0.0006840335,0.3648675,0.00034344956,0.0003156269,0.0017409484,0.0017374204,0.5333784,0.070207074,0.01080926,0.015608053,0.00019393465],"about_ca_topic_score_codex":0.0045412397,"about_ca_topic_score_gemma":0.006425088,"teacher_disagreement_score":0.008790963,"about_ca_system_score_codex":0.00049830356,"about_ca_system_score_gemma":0.0012679826,"threshold_uncertainty_score":0.010653853},"labels":[],"label_agreement":null},{"id":"W4241317238","doi":"10.1109/iwsc.2013.6613035","title":"Refactoring clones: A new perspective","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; clone (Java method); Computer science; Perspective (graphical); Software engineering; Software evolution; Software; Programming language; Position paper; Software system; World Wide Web; Artificial intelligence; Software construction","score_opus":0.022349811985856182,"score_gpt":0.27549396339230303,"score_spread":0.25314415140644686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241317238","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009032967,0.17059769,0.53742355,0.19733416,0.005925122,0.0000815161,0.00012927256,0.0006222016,0.07885348],"genre_scores_gemma":[0.32538924,0.24375717,0.3608977,0.023812216,0.020026363,0.00030007891,0.0002254207,0.0007010075,0.024890784],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9860844,0.0069099385,0.0008222398,0.0019497746,0.0036992182,0.0005344648],"domain_scores_gemma":[0.9405574,0.044729866,0.0021390792,0.0050126947,0.0059597376,0.0016012072],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01974039,0.0021351948,0.0027146318,0.00809584,0.0037718548,0.016907692,0.004729738,0.010700979,0.0053995815],"category_scores_gemma":[0.020906791,0.0011944234,0.0018120911,0.006120025,0.027517598,0.0501604,0.0062915566,0.0136891175,0.001466818],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034413388,0.00005228921,0.00032415215,0.0005698093,0.000024637544,0.0002655706,0.0027530435,0.00070895284,0.0005585096,0.94421244,0.0037947313,0.046701543],"study_design_scores_gemma":[0.000036367797,0.0001401786,0.0002765093,0.00067579426,0.00004523448,0.00077742507,0.0029473796,0.0039798757,0.00083184626,0.8378403,0.15238006,0.00006912019],"about_ca_topic_score_codex":0.0026877157,"about_ca_topic_score_gemma":0.0027916552,"teacher_disagreement_score":0.01974039,"about_ca_system_score_codex":0.004964612,"about_ca_system_score_gemma":0.004290214,"threshold_uncertainty_score":0.10439837},"labels":[],"label_agreement":null},{"id":"W4241558139","doi":"10.1145/1095430.1081711","title":"Automatic generation of suggestions for program investigation","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Source code; Intuition; Task (project management); Program analysis; Set (abstract data type); Dependency (UML); Static program analysis; Fuzzy logic; Software engineering; Code (set theory); Empirical research; Programming language; Software; Software development; Artificial intelligence; Systems engineering; Engineering","score_opus":0.043518429546045505,"score_gpt":0.28995529753509364,"score_spread":0.24643686798904813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241558139","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10954541,0.0006522758,0.82459885,0.0009603861,0.00020296978,0.0016260726,0.0017974334,0.05362362,0.006992965],"genre_scores_gemma":[0.14102663,0.0002018065,0.85156506,0.00011559059,0.00008406315,0.00072522106,0.00199723,0.0012808422,0.0030036194],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957217,0.0014025944,0.00023742448,0.0007809648,0.0016507738,0.00020652065],"domain_scores_gemma":[0.97180825,0.016998077,0.0020900476,0.0020178172,0.0064981133,0.0005877298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003392995,0.0022874398,0.001515992,0.0075449524,0.001104101,0.0014972555,0.002524392,0.0018750927,0.00790067],"category_scores_gemma":[0.03411709,0.0010077682,0.0011014033,0.0019179217,0.000620705,0.001907781,0.0014634876,0.0014696766,0.0029634417],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012184033,0.0006019251,0.011555568,0.001354768,0.00012745908,0.0011949124,0.0026476395,0.0135873975,0.062120155,0.008078161,0.03449635,0.86301726],"study_design_scores_gemma":[0.0004146569,0.0008414644,0.008774329,0.00048430113,0.00032719865,0.001182528,0.0019548128,0.8259107,0.09347965,0.015223586,0.051202957,0.00020379174],"about_ca_topic_score_codex":0.0021154387,"about_ca_topic_score_gemma":0.004831127,"teacher_disagreement_score":0.00790067,"about_ca_system_score_codex":0.00091612164,"about_ca_system_score_gemma":0.0027642485,"threshold_uncertainty_score":0.026430428},"labels":[],"label_agreement":null},{"id":"W4241904514","doi":"10.1007/978-3-540-73589-2_12","title":"Non-null References by Default in Java: Alleviating the Nullity Annotation Burden","year":2007,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Java; Programming language; Null (SQL); Annotation; Semantics (computer science); Java Modeling Language; Software engineering; Information retrieval; Natural language processing; Artificial intelligence; Java annotation; Database; Real time Java","score_opus":0.027390508133668202,"score_gpt":0.2866930219036231,"score_spread":0.2593025137699549,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241904514","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026609583,0.0008832961,0.909383,0.002317382,0.00095789577,0.00016850705,0.00047840993,0.04415064,0.015051304],"genre_scores_gemma":[0.31061935,0.0013735411,0.5935778,0.0033654533,0.0010570778,0.00032102494,0.001538615,0.06153409,0.026613047],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98066276,0.00611266,0.0023337717,0.0025448352,0.006726854,0.0016190419],"domain_scores_gemma":[0.8936813,0.04203352,0.0046487115,0.047146086,0.011009555,0.0014809539],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016284708,0.0020957612,0.0033334345,0.0026578659,0.0036243778,0.008013752,0.011775562,0.0054263333,0.010921168],"category_scores_gemma":[0.06668647,0.004854318,0.002392335,0.0037395465,0.007014562,0.023106152,0.018130368,0.010292818,0.0075890296],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014417239,0.00053229346,0.009164294,0.0019328584,0.0001644894,0.0020388651,0.0063059325,0.0061624115,0.03492986,0.3986762,0.07428639,0.4643648],"study_design_scores_gemma":[0.00043007106,0.00028080033,0.0017006041,0.0010700269,0.00077765825,0.0027500084,0.00122519,0.09303285,0.092632376,0.5684722,0.23702209,0.0006060421],"about_ca_topic_score_codex":0.0020493641,"about_ca_topic_score_gemma":0.0036643718,"teacher_disagreement_score":0.016284708,"about_ca_system_score_codex":0.0014148961,"about_ca_system_score_gemma":0.006200857,"threshold_uncertainty_score":0.08612275},"labels":[],"label_agreement":null},{"id":"W4241907497","doi":"10.1007/978-3-642-22655-7_5","title":"Using Structure-Based Recommendations to Facilitate Discoverability in APIs","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Discoverability; Computer science; Task (project management); Application programming interface; World Wide Web; Key (lock); Software engineering; Programming language; Engineering; Operating system","score_opus":0.08450920162949963,"score_gpt":0.299782116642758,"score_spread":0.21527291501325838,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241907497","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23522729,0.0026290184,0.69826716,0.0042586583,0.0006814795,0.0010817483,0.005020233,0.0294434,0.023390986],"genre_scores_gemma":[0.4426244,0.0010814955,0.5364292,0.00045518947,0.00020451512,0.00032069502,0.0077226465,0.0013802263,0.009781639],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9960278,0.0007831653,0.00033125872,0.00065754884,0.001974471,0.0002257304],"domain_scores_gemma":[0.98081434,0.011462048,0.0010269266,0.0034697568,0.0028226308,0.00040432895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027466828,0.0010717087,0.000990918,0.005062796,0.00129906,0.003230061,0.0020103133,0.0021632533,0.0061177746],"category_scores_gemma":[0.034149177,0.0008393528,0.0015501627,0.0029267778,0.0004755088,0.007822811,0.0020163092,0.0024547135,0.0027090774],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008017579,0.0010294504,0.031199606,0.00070103107,0.0003198058,0.0006426201,0.0018529049,0.016255552,0.016189743,0.018092003,0.03848489,0.8744307],"study_design_scores_gemma":[0.00023908193,0.00048543338,0.0135574555,0.0006138255,0.00069265085,0.0009224038,0.0013742518,0.80630976,0.04643638,0.05898699,0.07014996,0.00023184286],"about_ca_topic_score_codex":0.0104036005,"about_ca_topic_score_gemma":0.028890284,"teacher_disagreement_score":0.0104036005,"about_ca_system_score_codex":0.00095137017,"about_ca_system_score_gemma":0.0018587648,"threshold_uncertainty_score":0.02068609},"labels":[],"label_agreement":null},{"id":"W4241910944","doi":"10.1109/icsme.2014.19","title":"[Panels: 2 abstracts]","year":2014,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Alberta; University of British Columbia","funders":"","keywords":"Technical debt; Stakeholder; Bridging (networking); Metaphor; Computer science; Architecture; Debt; Knowledge management; Software; Business; Process management; Software development; Public relations; Political science; Computer security; Finance","score_opus":0.01884353693138684,"score_gpt":0.25706974933061694,"score_spread":0.2382262123992301,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241910944","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006694853,0.0040383334,0.0023023754,0.052159414,0.27523315,0.002994366,0.018788127,0.0034263642,0.64038837],"genre_scores_gemma":[0.0046816827,0.0029286146,0.0014127066,0.027766416,0.06084465,0.0021935846,0.0116428705,0.0019915784,0.88653785],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977914,0.00026460364,0.0002223726,0.0003471816,0.0010189231,0.0003554302],"domain_scores_gemma":[0.98572564,0.001504052,0.00062017597,0.0007190899,0.008712517,0.0027185753],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0027007035,0.0015464873,0.0019103241,0.0032556555,0.0041279104,0.007798289,0.0027338255,0.0063464437,0.774321],"category_scores_gemma":[0.013673071,0.0009818457,0.0019632052,0.0033299667,0.000621596,0.0036410587,0.0030374439,0.0049156393,0.5644704],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000032261083,0.0000069125076,0.00000957561,0.00006501958,0.0000015626357,0.000012860813,0.0000051539328,0.0000069268413,0.00007131033,0.00019912144,0.99528885,0.0043004714],"study_design_scores_gemma":[0.000049354105,0.000018898147,0.00026732584,0.00016016308,0.0000046622895,0.000026444748,0.000032716278,0.00004201619,0.00010913729,0.0006169148,0.99865913,0.000013205116],"about_ca_topic_score_codex":0.0036485097,"about_ca_topic_score_gemma":0.006108753,"teacher_disagreement_score":0.22567898,"about_ca_system_score_codex":0.0021670412,"about_ca_system_score_gemma":0.0037906345,"threshold_uncertainty_score":0.32190365},"labels":[],"label_agreement":null},{"id":"W4241947741","doi":"10.1109/icse.2003.1201219","title":"Hipikat: recommending pertinent software development artifacts","year":2003,"lang":"en","type":"article","venue":"25th International Conference on Software Engineering, 2003. Proceedings.","topic":"Software Engineering Research","field":"Computer Science","cited_by":173,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"U.S. Consumer Product Safety Commission","keywords":"Eclipse; Computer science; Task (project management); Context (archaeology); Open source; Open source software; Software; Software project management; Open-source software development; Software engineering; Software development; World Wide Web; Data science; Software construction; Engineering; Systems engineering; Programming language","score_opus":0.04131285254867295,"score_gpt":0.27387316265481426,"score_spread":0.2325603101061413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241947741","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25225955,0.0016736849,0.67373717,0.0027769292,0.00031565843,0.0030598736,0.0032401008,0.047715317,0.015221653],"genre_scores_gemma":[0.15544999,0.0005962084,0.8318896,0.00019931265,0.00006514585,0.00054425246,0.003267194,0.00064515835,0.0073430724],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99799734,0.00059668993,0.00019822121,0.0004392025,0.0006793251,0.00008924515],"domain_scores_gemma":[0.9868666,0.0074642133,0.0013094225,0.001879244,0.0017407752,0.0007397554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038040045,0.0012592431,0.0006703684,0.0050392016,0.0011790089,0.0025690042,0.0022278067,0.0017400947,0.005219799],"category_scores_gemma":[0.02931864,0.00077954103,0.0005792489,0.002375835,0.0004518373,0.0035078363,0.0023316504,0.00097189227,0.002398761],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048436812,0.0007509786,0.026977042,0.0021070123,0.0001863315,0.0012494235,0.008646168,0.0058830646,0.017359829,0.0030093796,0.037430875,0.89591557],"study_design_scores_gemma":[0.00087104493,0.0029420739,0.086575925,0.0019904221,0.0013533939,0.0060035703,0.019100754,0.25923538,0.07066793,0.028923485,0.5213999,0.00093613076],"about_ca_topic_score_codex":0.0046853297,"about_ca_topic_score_gemma":0.017569104,"teacher_disagreement_score":0.005219799,"about_ca_system_score_codex":0.0006533288,"about_ca_system_score_gemma":0.002566259,"threshold_uncertainty_score":0.0201177},"labels":[],"label_agreement":null},{"id":"W4242480381","doi":"10.4018/978-1-60566-418-7.ch020","title":"Modeling Defects in E-Projects","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software; Context (archaeology); Software development; Software engineering; Software development process; Field (mathematics); Process (computing); Systems engineering; Engineering; Geography","score_opus":0.03064925847424767,"score_gpt":0.2656258514144489,"score_spread":0.2349765929402012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4242480381","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40574604,0.0009921343,0.56958145,0.00083212595,0.00005180837,0.00021565176,0.00065642403,0.00058611674,0.0213383],"genre_scores_gemma":[0.9382882,0.0008817314,0.04840717,0.00006745003,0.000027606275,0.00028611408,0.0004456258,0.0000957833,0.0115004],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921596,0.00028622017,0.000052552226,0.0001726831,0.0001610113,0.00011155926],"domain_scores_gemma":[0.9951675,0.0033459482,0.00076027814,0.00023391396,0.00035206185,0.00014022388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015647061,0.00078765315,0.00057460595,0.0014092509,0.00030512747,0.0017765666,0.0021092019,0.0021464285,0.002803552],"category_scores_gemma":[0.006347899,0.00048399114,0.00085840275,0.0013668993,0.0009815823,0.002262697,0.0011839658,0.0009888037,0.0005126076],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000033661257,0.00015097125,0.007996513,0.00009154137,0.00003239775,0.0001756109,0.00028280885,0.8878666,0.0007366872,0.083285786,0.00073949643,0.018607898],"study_design_scores_gemma":[0.000011033857,0.000044535278,0.0015870384,0.000019337285,0.000015306097,0.000050802908,0.00007914046,0.97005635,0.00020598149,0.026784811,0.0011354882,0.000010154448],"about_ca_topic_score_codex":0.0076158424,"about_ca_topic_score_gemma":0.0045678085,"teacher_disagreement_score":0.0076158424,"about_ca_system_score_codex":0.0013050106,"about_ca_system_score_gemma":0.0008634421,"threshold_uncertainty_score":0.015143037},"labels":[],"label_agreement":null},{"id":"W4242528190","doi":"10.1145/568834.568837","title":"Recovering software requirements from system-user interaction traces","year":2002,"lang":"en","type":"article","venue":"Proceedings of the 14th international conference on Software engineering and knowledge engineering - SEKE '02","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Business process reengineering; Documentation; User requirements document; Software engineering; Task (project management); User interface; Software system; Process (computing); Software; Software development; Software requirements specification; Software design; Systems engineering; Programming language; Engineering","score_opus":0.03839105615279318,"score_gpt":0.2601434606518134,"score_spread":0.2217524044990202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4242528190","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5688795,0.00026838857,0.41935462,0.00034253646,0.00002645522,0.00039695116,0.003272108,0.005569014,0.0018904641],"genre_scores_gemma":[0.7706219,0.00023035979,0.22019072,0.000043744607,0.000014996211,0.0003180525,0.0072385347,0.00034549832,0.0009961756],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9960963,0.0010956331,0.00037543316,0.00041771104,0.0018133185,0.00020149806],"domain_scores_gemma":[0.97219837,0.01396583,0.0023053093,0.005540689,0.005494209,0.00049560145],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025911715,0.0012488128,0.00094484864,0.0046945405,0.00046947587,0.0013056664,0.0014700742,0.0013353563,0.00075982476],"category_scores_gemma":[0.03770504,0.00080700207,0.0008342318,0.0022353344,0.00040639928,0.001969357,0.0015197311,0.0014666371,0.0009260442],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092340936,0.0010195965,0.12628764,0.0009872618,0.00024425492,0.0022734918,0.0031158915,0.13162142,0.043215856,0.0071624983,0.005121727,0.6780271],"study_design_scores_gemma":[0.00005367227,0.0003095659,0.028432434,0.00008328611,0.00005509503,0.00067225186,0.00096985936,0.9304691,0.02396968,0.010245193,0.004674446,0.0000653757],"about_ca_topic_score_codex":0.007400215,"about_ca_topic_score_gemma":0.00901253,"teacher_disagreement_score":0.007400215,"about_ca_system_score_codex":0.0009780611,"about_ca_system_score_gemma":0.0017423226,"threshold_uncertainty_score":0.014714301},"labels":[],"label_agreement":null},{"id":"W4242845477","doi":"10.1145/2088883.2088889","title":"Report from the 2nd international workshop on replication in empirical software engineering research (RESER 2011)","year":2012,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Replication (statistics); Session (web analytics); Computer science; Empirical research; Joint (building); Software; Software engineering; Engineering; World Wide Web; Civil engineering; Biology","score_opus":0.08854124489477622,"score_gpt":0.3640275285554867,"score_spread":0.2754862836607105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4242845477","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007516914,0.057581782,0.3132458,0.3430378,0.12683782,0.0059780707,0.010799923,0.0071568135,0.12784517],"genre_scores_gemma":[0.0689668,0.061316375,0.3737647,0.07943733,0.038568042,0.014198071,0.047894098,0.015019697,0.3008349],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9138214,0.04641045,0.0059407875,0.0077210884,0.022380607,0.0037256218],"domain_scores_gemma":[0.74448717,0.09064354,0.0074056294,0.050332833,0.09206107,0.015069797],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.16028973,0.0025720126,0.002456745,0.0052530076,0.0050511644,0.015752945,0.00562503,0.009053925,0.072331324],"category_scores_gemma":[0.24415223,0.002335417,0.0040521454,0.004105371,0.0033071234,0.020935722,0.016819166,0.013487257,0.041869108],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029555234,0.00022792192,0.00095222617,0.00069808913,0.000080114696,0.00016130578,0.001695413,0.00067301217,0.0010077546,0.019339973,0.82002306,0.15484555],"study_design_scores_gemma":[0.000105120824,0.00013684135,0.0018551106,0.0013826413,0.00007071033,0.0001813099,0.0011373243,0.000940079,0.001569878,0.02649743,0.96598345,0.00014002327],"about_ca_topic_score_codex":0.0069950684,"about_ca_topic_score_gemma":0.007459397,"teacher_disagreement_score":0.83971024,"about_ca_system_score_codex":0.00535357,"about_ca_system_score_gemma":0.013326151,"threshold_uncertainty_score":0.8477033},"labels":[],"label_agreement":null},{"id":"W4243089840","doi":"10.1109/scam.2007.4362914","title":"A Framework for Studying Clones In Large Software Systems","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"clone (Java method); Computer science; Cloning (programming); Software system; Linux kernel; Source code; Software; Software maintenance; Software engineering; Software framework; Software development; Programming language; Operating system; Data mining; Software construction; Biology","score_opus":0.03959469923058079,"score_gpt":0.3274427764674381,"score_spread":0.2878480772368573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243089840","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050594592,0.0005238398,0.99120504,0.00037005622,0.0000152339835,0.00030799757,0.0005443123,0.0012485666,0.00072550087],"genre_scores_gemma":[0.02811986,0.00028007978,0.9698464,0.000050394097,0.00002204654,0.00046822437,0.000898459,0.00007543524,0.00023905578],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.992152,0.0028648805,0.0010877475,0.0015683173,0.0020066777,0.0003205134],"domain_scores_gemma":[0.9734624,0.017562965,0.0029052615,0.0034602154,0.0018722729,0.0007369665],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010381987,0.0022538097,0.0018376093,0.013716455,0.002235279,0.005596073,0.0043705665,0.0023251914,0.0022368385],"category_scores_gemma":[0.029984336,0.0014314082,0.0044776737,0.012145504,0.0037976147,0.010898092,0.004218329,0.0027219506,0.0006304363],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024792808,0.00062138215,0.02874975,0.0026988739,0.0007374311,0.001987382,0.010165509,0.09125344,0.015025089,0.46939275,0.010278074,0.36884233],"study_design_scores_gemma":[0.00011669412,0.00032417878,0.009957069,0.00053624855,0.00023305962,0.0018325022,0.0027596618,0.30457777,0.0063500647,0.6121533,0.060927648,0.00023188173],"about_ca_topic_score_codex":0.00878605,"about_ca_topic_score_gemma":0.007100902,"teacher_disagreement_score":0.013716455,"about_ca_system_score_codex":0.002322398,"about_ca_system_score_gemma":0.0030386425,"threshold_uncertainty_score":0.05490583},"labels":[],"label_agreement":null},{"id":"W4243181251","doi":"10.1109/icse.1998.671103","title":"Conceptual module querying for software reengineering","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Business process reengineering; Computer science; Software engineering; Source code; Set (abstract data type); Software; Conceptual model; Systems engineering; Database; Programming language; Engineering","score_opus":0.041729580875542564,"score_gpt":0.2513171246995899,"score_spread":0.20958754382404732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243181251","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022283357,0.00015967221,0.98908454,0.00021076632,0.000016933798,0.00009489859,0.00012300388,0.0066688904,0.0014129867],"genre_scores_gemma":[0.047097284,0.00029255226,0.9479704,0.00024277596,0.000052501677,0.00025202584,0.00096488,0.0017451852,0.001382538],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9924614,0.0033391037,0.0006984908,0.00096799963,0.0022164842,0.00031654164],"domain_scores_gemma":[0.98996514,0.0048228293,0.0005459486,0.003125107,0.0012729835,0.00026801316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008354731,0.0010308066,0.0012422274,0.0030828114,0.0012435883,0.005294846,0.0045572687,0.002567675,0.0065610725],"category_scores_gemma":[0.019847533,0.0012801125,0.0020684386,0.0039114314,0.0022486022,0.012898805,0.0045850934,0.0027803143,0.0024547684],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029413166,0.00015022676,0.0026868624,0.0008888095,0.00013001103,0.0005821392,0.002528102,0.011166859,0.014207344,0.5987314,0.023825092,0.34480906],"study_design_scores_gemma":[0.0001315441,0.00019272606,0.0011706646,0.00039001522,0.00016946529,0.0015571134,0.00055424013,0.2946254,0.026753573,0.40884304,0.26538575,0.00022647678],"about_ca_topic_score_codex":0.0019676276,"about_ca_topic_score_gemma":0.0023154598,"teacher_disagreement_score":0.008354731,"about_ca_system_score_codex":0.0014946545,"about_ca_system_score_gemma":0.0019041662,"threshold_uncertainty_score":0.044184625},"labels":[],"label_agreement":null},{"id":"W4243472257","doi":"10.1109/icse.2015.212","title":"Leveraging Informal Documentation to Summarize Classes and Methods in Context","year":2015,"lang":"en","type":"article","venue":"2015 IEEE/ACM 37th IEEE International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Documentation; Automatic summarization; Software documentation; Identifier; Context (archaeology); Program comprehension; Source code; Code (set theory); Software; Focus (optics); Software engineering; Task (project management); Benchmark (surveying); Internal documentation; Information retrieval; Software bug; World Wide Web; Software development; Programming language; Software system; Software construction; Set (abstract data type); Engineering; Systems engineering","score_opus":0.08999989506615488,"score_gpt":0.3885358259163932,"score_spread":0.29853593085023833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243472257","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20344482,0.0044601033,0.7524018,0.0020037054,0.0005024838,0.0013899071,0.010605476,0.014293587,0.010897964],"genre_scores_gemma":[0.33554053,0.0017189805,0.6358619,0.00033831323,0.00028050254,0.00096013496,0.01896906,0.0013288243,0.0050017717],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99461925,0.0022479708,0.0007893115,0.0009774729,0.0011997676,0.00016624853],"domain_scores_gemma":[0.9474647,0.026085125,0.0075226957,0.0072586387,0.011071212,0.00059754896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065187677,0.0016519986,0.0008122757,0.012295776,0.001403715,0.003493129,0.001139884,0.001258197,0.0024177295],"category_scores_gemma":[0.063877985,0.0006163729,0.00046938463,0.0056281406,0.0005807752,0.0049265013,0.0037996704,0.0014226922,0.0014995671],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046925235,0.000208945,0.023770327,0.0035866373,0.00014772512,0.00063843577,0.025600953,0.0063298587,0.023842644,0.009541722,0.026090307,0.87977326],"study_design_scores_gemma":[0.00030779786,0.0013861906,0.06986603,0.004556528,0.0009782981,0.0019445153,0.022847736,0.22140616,0.08065445,0.07275086,0.5227016,0.00059978565],"about_ca_topic_score_codex":0.0042999843,"about_ca_topic_score_gemma":0.008426572,"teacher_disagreement_score":0.012295776,"about_ca_system_score_codex":0.0009088689,"about_ca_system_score_gemma":0.0023238335,"threshold_uncertainty_score":0.03447497},"labels":[],"label_agreement":null},{"id":"W4243774012","doi":"10.1145/2666357.2597823","title":"em-SPADE","year":2014,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Compiler; Programming language; Preprocessor; Software; Parallel computing","score_opus":0.01955025228379702,"score_gpt":0.2625397247484502,"score_spread":0.2429894724646532,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243774012","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0354205,0.00071286067,0.5727072,0.00062776776,0.00043936798,0.00035934322,0.009575757,0.33662394,0.04353335],"genre_scores_gemma":[0.23640779,0.00059991947,0.6166499,0.0015468734,0.00013590336,0.0007119187,0.039616875,0.045960393,0.058370475],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99878293,0.00019900216,0.00012138058,0.0003083414,0.0004790133,0.00010936818],"domain_scores_gemma":[0.9971847,0.0007112329,0.00021227224,0.0011973183,0.0006271534,0.000067254325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001104877,0.00092822115,0.00053041853,0.0008490494,0.00035426187,0.0011873591,0.0016984481,0.0007012492,0.020477543],"category_scores_gemma":[0.00464315,0.0009985099,0.00097016036,0.00049704715,0.00051782595,0.0026905406,0.001966086,0.0014360489,0.013756786],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019208579,0.0004660737,0.0113338325,0.0013771497,0.00027169264,0.0006519591,0.00056395086,0.021702264,0.04199682,0.051537532,0.35407683,0.514101],"study_design_scores_gemma":[0.00039752084,0.0005631821,0.003328734,0.00018333827,0.00015249092,0.0013484017,0.00008958051,0.17743874,0.15597649,0.028894639,0.6315184,0.00010845085],"about_ca_topic_score_codex":0.00069925253,"about_ca_topic_score_gemma":0.0012353967,"teacher_disagreement_score":0.020477543,"about_ca_system_score_codex":0.00041380024,"about_ca_system_score_gemma":0.0011776591,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4243809191","doi":"10.4018/978-1-4666-3679-8.ch012","title":"Software Security Engineering – Part I","year":2013,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; University of Victoria","funders":"","keywords":"Software security assurance; Software development; Social software engineering; Computer science; Personal software process; Software peer review; Software construction; Security bug; Software development process; Security engineering; Software engineering; Package development process; Software; Computer security; Security service; Information security; Operating system","score_opus":0.014445339594080387,"score_gpt":0.2282020759556463,"score_spread":0.2137567363615659,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243809191","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004782842,0.41357997,0.096388236,0.0065514157,0.005273207,0.0002795209,0.00039739642,0.00085062726,0.47189683],"genre_scores_gemma":[0.05280198,0.44319224,0.058446772,0.0044204486,0.0046685454,0.00041920834,0.0015200091,0.0008293398,0.43370146],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99923134,0.00013178209,0.0000687323,0.00015682678,0.00036591996,0.000045436467],"domain_scores_gemma":[0.9994155,0.00025438247,0.000041750485,0.000098618104,0.00014710487,0.000042645257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00062138453,0.0012568772,0.00064346753,0.0022634151,0.0007266433,0.0030422518,0.000736763,0.0014795095,0.01736208],"category_scores_gemma":[0.0012885886,0.0006085827,0.0006833861,0.0030181834,0.0017100192,0.0033005676,0.0014647201,0.0025403183,0.0119469445],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020453483,0.00014603825,0.00038137246,0.0020188354,0.0000313289,0.0001778751,0.0012300797,0.0032881491,0.0033720515,0.18847395,0.13686754,0.66399235],"study_design_scores_gemma":[0.000003576788,0.00004975182,0.00061215484,0.0012134376,0.0000072528214,0.00050952146,0.00015329156,0.00083026954,0.0007858237,0.07940859,0.91641057,0.0000157993],"about_ca_topic_score_codex":0.0008240085,"about_ca_topic_score_gemma":0.0008594192,"teacher_disagreement_score":0.01736208,"about_ca_system_score_codex":0.0013476876,"about_ca_system_score_gemma":0.0014041729,"threshold_uncertainty_score":0.058081985},"labels":[],"label_agreement":null},{"id":"W4244354649","doi":"10.4018/9781605660608.ch055","title":"Constructivist Learning During Software Development","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Laurentian University","funders":"","keywords":"Programmer; Constructivist teaching methods; Computer science; Denial; Documentation; Process (computing); Zone of proximal development; Development (topology); Cognitive science; Knowledge management; Mathematics education; Software engineering; Psychology; Programming language; Teaching method; Mathematics","score_opus":0.018925662165124438,"score_gpt":0.22807582701891668,"score_spread":0.20915016485379223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244354649","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05644147,0.035322107,0.36389053,0.020629525,0.00044372174,0.00016367446,0.0000663267,0.00037630353,0.52266634],"genre_scores_gemma":[0.73532397,0.024744067,0.14109676,0.0017353452,0.0002985677,0.00034875888,0.00012505738,0.00023880313,0.09608867],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.99806434,0.0012059129,0.000040305516,0.00017819893,0.00041862077,0.00009269496],"domain_scores_gemma":[0.9968232,0.0027373852,0.000081824255,0.00019886946,0.00008936673,0.00006934097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028422677,0.00054387335,0.00026843036,0.00074317754,0.001076495,0.0035134626,0.0012655557,0.0015429137,0.0032481828],"category_scores_gemma":[0.004247232,0.00036166538,0.00027485428,0.0009927001,0.012990291,0.0054344535,0.0025672026,0.0040866802,0.00069163746],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000018434665,0.00004713764,0.0004176309,0.00029644984,0.000008115497,0.00016361104,0.026492476,0.0033875238,0.0005464347,0.8712212,0.0039182594,0.093482755],"study_design_scores_gemma":[0.000027882945,0.000051560884,0.0005169315,0.0005397107,0.000008301196,0.00025064347,0.0031830864,0.004719127,0.0012010183,0.7949419,0.19454065,0.000019122283],"about_ca_topic_score_codex":0.0017822712,"about_ca_topic_score_gemma":0.0029390918,"teacher_disagreement_score":0.004495026,"about_ca_system_score_codex":0.004495026,"about_ca_system_score_gemma":0.0024055268,"threshold_uncertainty_score":0.032613814},"labels":[],"label_agreement":null},{"id":"W4244478245","doi":"10.1115/1.4037817","title":"Automated Extraction of Function Knowledge From Text","year":2017,"lang":"en","type":"article","venue":"Journal of Mechanical Design","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Autodesk (Canada)","funders":"Oregon State University","keywords":"WordNet; Computer science; Natural language processing; Word2vec; Knowledge base; Parsing; Function (biology); Artificial intelligence; Information retrieval; Knowledge extraction; Artifact (error); Information extraction","score_opus":0.059221493670442445,"score_gpt":0.329604316374932,"score_spread":0.27038282270448954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244478245","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.088674136,0.0037419894,0.81136054,0.0012914259,0.0002804873,0.0012687667,0.044509813,0.029783886,0.01908896],"genre_scores_gemma":[0.15739004,0.00203679,0.75531083,0.00023857549,0.00016556292,0.00075405865,0.0778993,0.0011985311,0.0050063566],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99813914,0.00034634658,0.00026374465,0.00050198147,0.0006488894,0.00009991533],"domain_scores_gemma":[0.99207497,0.004350832,0.00085359917,0.00085569266,0.001751268,0.0001136049],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012968496,0.0016970494,0.0009882047,0.01753793,0.00097501173,0.002200155,0.0014091422,0.0009943339,0.004141623],"category_scores_gemma":[0.009349975,0.00061252073,0.0012499826,0.0072302134,0.00072861434,0.004731256,0.0015998876,0.0010000165,0.003672172],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015529511,0.00022337577,0.005551849,0.0024035678,0.00012778747,0.0018241496,0.0014818314,0.0041726143,0.03141342,0.0093550095,0.031441387,0.91184974],"study_design_scores_gemma":[0.00016050122,0.00036713073,0.035179082,0.0024165804,0.0006844908,0.0048282337,0.0046130596,0.25255957,0.18144283,0.085082434,0.43228987,0.00037620522],"about_ca_topic_score_codex":0.0032486552,"about_ca_topic_score_gemma":0.0046350104,"teacher_disagreement_score":0.01753793,"about_ca_system_score_codex":0.0010138038,"about_ca_system_score_gemma":0.0024886902,"threshold_uncertainty_score":0.0138551},"labels":[],"label_agreement":null},{"id":"W4244764614","doi":"10.1109/wetsom.2017.7","title":"A Heuristic for Estimating the Impact of Lingering Defects: Can Debt Analogy Be Used as a Metric?","year":2017,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Technical debt; Computer science; Metric (unit); Software; Debt; Heuristic; Software bug; Software quality; Reliability engineering; Software development; Finance; Business; Artificial intelligence; Operations management; Engineering; Operating system","score_opus":0.047002322771078006,"score_gpt":0.3645933698870152,"score_spread":0.31759104711593716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244764614","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6418853,0.0053251465,0.33856934,0.0020985564,0.00015602533,0.0007948119,0.0037137954,0.0025052237,0.0049518896],"genre_scores_gemma":[0.86022747,0.00030628228,0.13623846,0.00017487773,0.000095847485,0.00021557542,0.0019905,0.00005961644,0.00069128425],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9976554,0.0009102012,0.00026641253,0.00050862384,0.00040168213,0.00025772094],"domain_scores_gemma":[0.9791311,0.016065098,0.002053531,0.00069267355,0.0015108562,0.0005468322],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00416794,0.0013845369,0.0014431758,0.007516299,0.00057321,0.0022995207,0.0016615521,0.002386888,0.0011323305],"category_scores_gemma":[0.016772518,0.00046873,0.000671762,0.0041983426,0.00087197527,0.0021081816,0.00062520156,0.0010454357,0.0003458401],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083944516,0.0013885577,0.19702904,0.00087163877,0.0005783199,0.0003759037,0.00036352917,0.47429544,0.004084087,0.004347536,0.01279123,0.3030353],"study_design_scores_gemma":[0.000084481995,0.00030852287,0.01544085,0.000060117432,0.00008013169,0.00020064902,0.000162553,0.97654164,0.0013696464,0.004578462,0.0011310925,0.000041840372],"about_ca_topic_score_codex":0.004992184,"about_ca_topic_score_gemma":0.0070995237,"teacher_disagreement_score":0.007516299,"about_ca_system_score_codex":0.0020297088,"about_ca_system_score_gemma":0.0023011686,"threshold_uncertainty_score":0.022042453},"labels":[],"label_agreement":null},{"id":"W4244949983","doi":"10.1109/icse.2000.870428","title":"A replicated assessment and comparison of common software cost modeling techniques","year":2002,"lang":"en","type":"article","venue":"Proceedings of the 2000 International Conference on Software Engineering. ICSE 2000 the New Millennium","topic":"Software Engineering Research","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Reputation; Software quality; Software; Quality (philosophy); Risk analysis (engineering); Software quality control; Software metric; Order (exchange); Product (mathematics); Software quality analyst; Software sizing; Reliability engineering; Software development; Software construction; Business; Engineering; Finance","score_opus":0.04642133386724595,"score_gpt":0.30996938636266536,"score_spread":0.2635480524954194,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244949983","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4376565,0.0075992844,0.5280461,0.0010990157,0.00024213904,0.0010158197,0.0011410813,0.0014083496,0.021791792],"genre_scores_gemma":[0.6559332,0.003094549,0.3374298,0.000075580705,0.00006883264,0.0005666256,0.0011130511,0.00025361302,0.0014647883],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97276974,0.013116113,0.0018543337,0.0013875265,0.010328926,0.0005432725],"domain_scores_gemma":[0.90782577,0.06021498,0.005141978,0.009200277,0.017168919,0.0004481029],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03139191,0.0018538074,0.0012310434,0.010835076,0.00097508,0.0034390972,0.0029924975,0.0019289568,0.0018301855],"category_scores_gemma":[0.10032503,0.0006162188,0.003052225,0.0075547807,0.0008546925,0.005081871,0.002459105,0.0017082195,0.00050807995],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012625402,0.0011384431,0.022565518,0.0012999866,0.00082368834,0.000106623644,0.0012070842,0.23557384,0.004424768,0.031741362,0.0026845078,0.69717157],"study_design_scores_gemma":[0.00014667603,0.0021848017,0.031730317,0.00053258694,0.00056906167,0.0002820947,0.0009953559,0.92741525,0.0067758732,0.019706044,0.00943227,0.00022957662],"about_ca_topic_score_codex":0.0064153303,"about_ca_topic_score_gemma":0.005845805,"teacher_disagreement_score":0.9686081,"about_ca_system_score_codex":0.0038584012,"about_ca_system_score_gemma":0.002190439,"threshold_uncertainty_score":0.16601825},"labels":[],"label_agreement":null},{"id":"W4244955996","doi":"10.1145/3393934.3278127","title":"Exploring feature interactions without specifications: a controlled experiment","year":2020,"lang":"en","type":"article","venue":"ACM SIGPLAN Notices","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Feature (linguistics); Graph; Control flow; Control flow graph; Data mining; Feature model; Information flow; Machine learning; Artificial intelligence; Theoretical computer science; Human–computer interaction; Programming language; Software","score_opus":0.22279030538800812,"score_gpt":0.3249284050983488,"score_spread":0.10213809971034069,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244955996","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9831548,0.00015315639,0.009894477,0.00013201761,0.00010456218,0.004561371,0.0005280465,0.00037026222,0.001101388],"genre_scores_gemma":[0.9272426,0.00024650226,0.050241187,0.00073396077,0.0001720822,0.015657958,0.0010162828,0.00020017952,0.004489172],"study_design_codex":"bench_or_experimental","study_design_gemma":"nonrandomized_trial","domain_scores_codex":[0.9958015,0.0020322185,0.00035252344,0.00097550434,0.00046862155,0.0003695612],"domain_scores_gemma":[0.92692024,0.06262294,0.003021704,0.0033694918,0.002299109,0.0017665165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007172884,0.0018308312,0.0010435983,0.000654164,0.0009473758,0.0012200599,0.002138994,0.002271961,0.009178252],"category_scores_gemma":[0.02321655,0.0009898461,0.0007416409,0.00041333423,0.0014031302,0.0019678795,0.0014676682,0.0016591462,0.0014500824],"study_design_candidate":"nonrandomized_trial","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.10167422,0.2906597,0.035138182,0.008502848,0.0008476301,0.0028437767,0.033471543,0.011738352,0.31095546,0.0030231397,0.010698731,0.19044638],"study_design_scores_gemma":[0.06696827,0.5799923,0.09508623,0.00078580965,0.0019852652,0.0018414983,0.010140975,0.050151687,0.15072566,0.010626125,0.030895548,0.00080062466],"about_ca_topic_score_codex":0.00063713663,"about_ca_topic_score_gemma":0.0008277664,"teacher_disagreement_score":0.009178252,"about_ca_system_score_codex":0.00037236177,"about_ca_system_score_gemma":0.00087239296,"threshold_uncertainty_score":0.037934303},"labels":[],"label_agreement":null},{"id":"W4245225971","doi":"10.1145/1140124.1140237","title":"Finite automata models for CS problem with binary semaphore","year":2006,"lang":"en","type":"article","venue":"Proceedings of the 11th annual SIGCSE conference on Innovation and technology in computer science education","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nipissing University","funders":"","keywords":"Semaphore; Computer science; Automaton; Synchronization (alternating current); Simplicity; Finite-state machine; Deterministic finite automaton; Theoretical computer science; Visualization; Parallelism (grammar); Binary number; Parallel computing; Algorithm; Programming language; Mathematics; Arithmetic; Artificial intelligence; Channel (broadcasting)","score_opus":0.016994059867617324,"score_gpt":0.263056050210883,"score_spread":0.2460619903432657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4245225971","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022479983,0.0008340241,0.9587699,0.0008022667,0.00014555609,0.0001484255,0.0003718362,0.00041282544,0.016035186],"genre_scores_gemma":[0.6110095,0.0016653095,0.35587516,0.0003410502,0.00030651674,0.00087390596,0.0009791787,0.0002085691,0.028740847],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99913836,0.00028388924,0.000063726286,0.00020898631,0.00019296266,0.000112085065],"domain_scores_gemma":[0.9979353,0.0014361024,0.00017926672,0.00017693285,0.00018384618,0.00008866969],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008172718,0.00086832885,0.0005920468,0.00074053876,0.00088678027,0.0019208888,0.00134183,0.0016020316,0.0060005574],"category_scores_gemma":[0.003484699,0.00035092534,0.0014719715,0.0008415162,0.0012597428,0.002355278,0.0011116973,0.0020050386,0.0008297248],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000066946035,0.00007893379,0.0005872029,0.00016059859,0.000030760642,0.00033040214,0.00026779884,0.22884949,0.0019110496,0.75253826,0.0024857514,0.01269278],"study_design_scores_gemma":[0.000019641488,0.000036650854,0.00012417154,0.000038080703,0.000024314399,0.00012356753,0.00008245656,0.60927945,0.0006701993,0.38227457,0.0073093898,0.000017513641],"about_ca_topic_score_codex":0.005543778,"about_ca_topic_score_gemma":0.006509818,"teacher_disagreement_score":0.0060005574,"about_ca_system_score_codex":0.001727665,"about_ca_system_score_gemma":0.0020719888,"threshold_uncertainty_score":0.020073831},"labels":[],"label_agreement":null},{"id":"W4245283101","doi":"10.1002/smr.395","title":"Special Issue on Search‐Based Software Maintenance","year":2008,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Search-based software engineering; Software engineering; Computer science; Software construction; Software maintenance; Software development; Software sizing; Code refactoring; Software; Reliability engineering; Systems engineering; Engineering; Programming language","score_opus":0.049937845687060084,"score_gpt":0.330768104488588,"score_spread":0.2808302588015279,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4245283101","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00079939497,0.061641295,0.0062365658,0.023170603,0.80984426,0.00021892707,0.0009790423,0.001038582,0.09607145],"genre_scores_gemma":[0.0027664558,0.047960367,0.001782775,0.008020706,0.8019628,0.00022043765,0.0016620007,0.00097500056,0.13464938],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974433,0.00029157774,0.00020089463,0.00062547525,0.0011818161,0.00025684226],"domain_scores_gemma":[0.99381495,0.0021571903,0.00047723085,0.0005678869,0.0018478696,0.0011348395],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016301358,0.002843429,0.0029632584,0.0035613787,0.0018986714,0.007375859,0.0033701872,0.0046803053,0.13908736],"category_scores_gemma":[0.0070444434,0.00080855115,0.0016013935,0.0028204648,0.0013156966,0.0054340824,0.0033649236,0.0053242473,0.06425486],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039339524,0.00006505654,0.00009999561,0.0003962968,0.000018359246,0.00006701164,0.000023014503,0.00024628648,0.000295123,0.0023974152,0.9525603,0.043791763],"study_design_scores_gemma":[0.00001221492,0.000055471926,0.0002590165,0.00028251976,0.000012700175,0.00012682934,0.000020423238,0.0002846737,0.0001334426,0.002500788,0.99629956,0.0000124320795],"about_ca_topic_score_codex":0.0008490119,"about_ca_topic_score_gemma":0.0013562914,"teacher_disagreement_score":0.13908736,"about_ca_system_score_codex":0.002014485,"about_ca_system_score_gemma":0.0020018837,"threshold_uncertainty_score":0.4652936},"labels":[],"label_agreement":null},{"id":"W4246180958","doi":"10.1109/icpc.2015.16","title":"Could We Infer Unordered API Usage Patterns Only Using the Library Source Code?","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Source code; Code (set theory); Programming language; Open source; Information retrieval; World Wide Web; Software; Set (abstract data type)","score_opus":0.06523818212866261,"score_gpt":0.2918464819780752,"score_spread":0.2266082998494126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4246180958","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54952794,0.0009357028,0.43441778,0.00084555394,0.0000447173,0.00026197662,0.0038301977,0.006186951,0.003949163],"genre_scores_gemma":[0.78265506,0.00064994517,0.20756055,0.00021919349,0.000033239998,0.00026269123,0.0061103534,0.0007499152,0.0017590442],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970626,0.00053085986,0.00028817958,0.0006891229,0.0011562107,0.00027286715],"domain_scores_gemma":[0.98351085,0.006486734,0.0030367197,0.0035801,0.0031010436,0.00028444984],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019694348,0.0010949157,0.0008679963,0.0048639243,0.00044884224,0.0013739694,0.0015761175,0.0009637461,0.0010681516],"category_scores_gemma":[0.021945111,0.000709917,0.0009596751,0.0055233343,0.0005682817,0.0040873634,0.0010456484,0.00095758133,0.001371251],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032588124,0.00040813236,0.39367658,0.0009651546,0.00042144945,0.0010474527,0.0015074661,0.008437273,0.022822404,0.0024845146,0.0047429777,0.5631607],"study_design_scores_gemma":[0.000080381775,0.000459381,0.295235,0.0005979009,0.0007459171,0.0039422717,0.0034866652,0.549549,0.075997755,0.037319195,0.032399457,0.00018712309],"about_ca_topic_score_codex":0.007144852,"about_ca_topic_score_gemma":0.0152004585,"teacher_disagreement_score":0.007144852,"about_ca_system_score_codex":0.00036457894,"about_ca_system_score_gemma":0.0013566144,"threshold_uncertainty_score":0.014206529},"labels":[],"label_agreement":null},{"id":"W4246265994","doi":"10.1145/1083292.1083302","title":"Quality, cleanroom and formal methods","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Cleanroom; Computer science; Software engineering; Software development; Quality (philosophy); Formal methods; Software quality; Systems engineering; Software; Engineering; Programming language","score_opus":0.05331196725728852,"score_gpt":0.41162451680411566,"score_spread":0.3583125495468271,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4246265994","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020055706,0.004514116,0.98119307,0.0040548234,0.0002454743,0.000074627715,0.000028986142,0.000502618,0.007380749],"genre_scores_gemma":[0.093317986,0.0061055785,0.8886523,0.0014228744,0.0006004014,0.00031069395,0.00013002769,0.00037432514,0.009085824],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9870803,0.004623295,0.0010639654,0.0014582194,0.005162316,0.00061200216],"domain_scores_gemma":[0.9706181,0.01784864,0.0021635303,0.004979526,0.0038659398,0.0005242132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014658874,0.001413262,0.0011911909,0.004224651,0.001547016,0.00493003,0.0035888182,0.002165118,0.0040978384],"category_scores_gemma":[0.023818867,0.0010492841,0.0027512764,0.0020092805,0.01640641,0.013004084,0.004527959,0.005701121,0.0009548649],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020469528,0.0000453552,0.00029209675,0.00043719725,0.000032794982,0.000088094064,0.0005088844,0.008706048,0.0007726272,0.92237765,0.0019941442,0.06472458],"study_design_scores_gemma":[0.00004742708,0.000075956195,0.00011581496,0.00029322974,0.000039183073,0.00019510486,0.00013284,0.015125993,0.0017313438,0.92310745,0.05907815,0.000057591704],"about_ca_topic_score_codex":0.0035881514,"about_ca_topic_score_gemma":0.0023079352,"teacher_disagreement_score":0.014658874,"about_ca_system_score_codex":0.0038116975,"about_ca_system_score_gemma":0.0041184085,"threshold_uncertainty_score":0.07752442},"labels":[],"label_agreement":null},{"id":"W4247181191","doi":"10.6028/nist.sp.500-283","title":"Report on the third static analysis tool exposition (SATE 2010)","year":2011,"lang":"en","type":"report","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"National Institute of Standards and Technology","keywords":"Exposition (narrative); Computer science; Art; Literature","score_opus":0.06399971136500837,"score_gpt":0.30598247709664633,"score_spread":0.24198276573163796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247181191","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.076355845,0.0034481185,0.26252744,0.03216472,0.010575973,0.007770154,0.16501912,0.051853653,0.3902849],"genre_scores_gemma":[0.06952862,0.0017910824,0.26390484,0.0058695297,0.00078644324,0.0045239814,0.24984069,0.013963397,0.38979137],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97909,0.0022867604,0.00069973327,0.0010517106,0.015554835,0.0013170062],"domain_scores_gemma":[0.94736904,0.0072540236,0.0012281624,0.006684409,0.034822915,0.002641526],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021721574,0.0013022821,0.0007983262,0.0053498666,0.00198315,0.0050067315,0.00196301,0.003141133,0.038888983],"category_scores_gemma":[0.0367045,0.0011534146,0.0014248908,0.0034164754,0.0006073765,0.0038705268,0.0032893773,0.0040594824,0.037188303],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030396323,0.0006369985,0.002944589,0.00021685765,0.000047995818,0.00027026774,0.0004109599,0.0018297759,0.011967495,0.0073823663,0.8103782,0.16361043],"study_design_scores_gemma":[0.00013856417,0.00049427873,0.0087705925,0.00023408775,0.000046880636,0.00035018224,0.00023572851,0.002821983,0.016899733,0.0020224198,0.9678674,0.00011814172],"about_ca_topic_score_codex":0.014218181,"about_ca_topic_score_gemma":0.018795503,"teacher_disagreement_score":0.038888983,"about_ca_system_score_codex":0.0030108015,"about_ca_system_score_gemma":0.0100254845,"threshold_uncertainty_score":0.13009661},"labels":[],"label_agreement":null},{"id":"W4247453041","doi":"10.1145/634636.586105","title":"Combining static and dynamic data in code visualization","year":2002,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Compiler; Programming language; Debugging; Visualization; Metadata; Software visualization; Source code; Interface (matter); Static analysis; Extensibility; Compile time; Software; Software development; Operating system; Component-based software engineering; Data mining","score_opus":0.03998723542697999,"score_gpt":0.29331550146482127,"score_spread":0.2533282660378413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247453041","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016571594,0.00072840095,0.9638457,0.0008967447,0.000068482324,0.00006347899,0.0003190541,0.012253021,0.0052534845],"genre_scores_gemma":[0.2017345,0.0011566847,0.7915814,0.00018909478,0.00008794906,0.00014600619,0.0006487404,0.0027074013,0.0017482272],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.997609,0.0008683232,0.00014123145,0.0002815403,0.00095174205,0.00014806123],"domain_scores_gemma":[0.993064,0.0032518,0.00047701987,0.0018374039,0.00107001,0.00029981893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004311469,0.0009626294,0.0009637681,0.0054113762,0.0010278901,0.005754801,0.0013375959,0.0011836478,0.0030549413],"category_scores_gemma":[0.012857463,0.0010352491,0.00081512984,0.00361376,0.0016644535,0.0068948437,0.004711504,0.002245316,0.00087872415],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005136569,0.00020580906,0.00782467,0.0008490892,0.00015653069,0.00071071275,0.0045070993,0.044985227,0.050263852,0.109151386,0.018180404,0.76265156],"study_design_scores_gemma":[0.00017544936,0.00029794234,0.010835709,0.00069355906,0.00027821594,0.0018057827,0.0020798373,0.41880697,0.10717062,0.27146333,0.18579589,0.0005966977],"about_ca_topic_score_codex":0.0026852626,"about_ca_topic_score_gemma":0.0031174007,"teacher_disagreement_score":0.005754801,"about_ca_system_score_codex":0.0006732172,"about_ca_system_score_gemma":0.0010182604,"threshold_uncertainty_score":0.022801518},"labels":[],"label_agreement":null},{"id":"W4247511377","doi":"10.1109/iwpse.2004.1334772","title":"Aiding comprehension of cloning through categorization","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"clone (Java method); Program comprehension; Computer science; Software maintenance; False positive paradox; Cloning (programming); Function (biology); Programming language; Software; Source code; Filter (signal processing); Set (abstract data type); Software system; Artificial intelligence; Biology; Genetics; Gene","score_opus":0.02673080414723087,"score_gpt":0.2748554369435063,"score_spread":0.24812463279627542,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247511377","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24316405,0.0008045321,0.7349719,0.0017409644,0.00008839491,0.0004220443,0.00055729056,0.01295031,0.005300499],"genre_scores_gemma":[0.47673947,0.0004648731,0.5176163,0.00047742348,0.000057202727,0.00025314925,0.001519362,0.0005485961,0.0023236983],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9897077,0.0047139167,0.0010597719,0.0016929511,0.002357809,0.00046791197],"domain_scores_gemma":[0.8984673,0.069763325,0.006296053,0.01123264,0.013220021,0.0010206052],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010597673,0.0011702714,0.0014335171,0.0073748957,0.00127172,0.005299308,0.0028666023,0.0028780212,0.0025868497],"category_scores_gemma":[0.08051569,0.00068578275,0.0008845769,0.0036852856,0.0013369475,0.012660286,0.003666606,0.0018164503,0.0011852386],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005756726,0.00047370995,0.04897824,0.0011249607,0.000106529,0.00075078243,0.024380976,0.00656057,0.065530956,0.02246205,0.0072318367,0.82182366],"study_design_scores_gemma":[0.0001735549,0.0014965885,0.0559945,0.0010887592,0.00043021954,0.0060340418,0.019288717,0.48563892,0.14925091,0.17217082,0.10781986,0.0006130856],"about_ca_topic_score_codex":0.0025065204,"about_ca_topic_score_gemma":0.0023981747,"teacher_disagreement_score":0.010597673,"about_ca_system_score_codex":0.0013941043,"about_ca_system_score_gemma":0.0023052457,"threshold_uncertainty_score":0.056046546},"labels":[],"label_agreement":null},{"id":"W4247693085","doi":"10.1109/icse.2001.919086","title":"On the syllogistic structure of object-oriented programming","year":2005,"lang":"en","type":"article","venue":"Proceedings of the 23rd International Conference on Software Engineering. ICSE 2001","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Syllogism; Computer science; Class hierarchy; Theoretical computer science; Hierarchy; Class (philosophy); Graph; Object-oriented programming; Programming language; Artificial intelligence; Epistemology","score_opus":0.021668304380192165,"score_gpt":0.2598943500563406,"score_spread":0.23822604567614844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247693085","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041572426,0.0052730087,0.78565663,0.014214128,0.0005386034,0.00009975769,0.00025300062,0.0007603882,0.15163212],"genre_scores_gemma":[0.7137669,0.0041385875,0.26197895,0.0023826864,0.0013699499,0.00038052333,0.00048691768,0.0003945211,0.015101028],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9968443,0.0013238554,0.00017789115,0.00039250974,0.0009963171,0.000265141],"domain_scores_gemma":[0.9936792,0.0040925457,0.00038331538,0.0006319127,0.0008541906,0.00035881266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030768071,0.00037862128,0.00047149402,0.0021143488,0.0034955486,0.005376226,0.000848006,0.0015598974,0.0053240554],"category_scores_gemma":[0.010235072,0.00069891976,0.0006776765,0.0018453484,0.012867042,0.010443583,0.003612974,0.0036303645,0.0009823752],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000039446695,0.000004692442,0.00010325194,0.000012033622,0.0000010836966,0.000016554266,0.00030484053,0.0003025346,0.00007095718,0.9940294,0.00059381756,0.0045568957],"study_design_scores_gemma":[0.000003458559,0.0000035197304,0.00009802027,0.000017612958,0.0000016293354,0.000024902884,0.000042545573,0.0018294443,0.00007304546,0.9898183,0.008081626,0.0000059623167],"about_ca_topic_score_codex":0.0053722397,"about_ca_topic_score_gemma":0.0046230494,"teacher_disagreement_score":0.005376226,"about_ca_system_score_codex":0.002851337,"about_ca_system_score_gemma":0.0024119066,"threshold_uncertainty_score":0.020687997},"labels":[],"label_agreement":null},{"id":"W4247706268","doi":"10.18293/seke2015-42","title":"Using peak analysis for identifying lagged effects between software metrics","year":2015,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Software; Operating system","score_opus":0.09751997057660587,"score_gpt":0.3208275290474762,"score_spread":0.2233075584708703,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247706268","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15962486,0.00061362924,0.834865,0.00016468186,0.00015137358,0.00014415296,0.00081273244,0.0018027071,0.0018209778],"genre_scores_gemma":[0.8193314,0.00028644744,0.17827696,0.000060656166,0.0001434209,0.0001459037,0.00081894343,0.00016877819,0.0007674794],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99866307,0.0003684484,0.00011067957,0.00043391358,0.00027788905,0.00014601847],"domain_scores_gemma":[0.98951614,0.0071933027,0.0009546273,0.001061178,0.0010013419,0.00027334623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041596517,0.0009480109,0.0007870057,0.004841893,0.0004994177,0.0014663796,0.0008205156,0.00063647504,0.0025232334],"category_scores_gemma":[0.013018144,0.00040504037,0.000989132,0.0031947948,0.00050844165,0.0022559445,0.0012340801,0.0010214662,0.00054104894],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021964544,0.0009579282,0.18430078,0.0009121266,0.0013514779,0.0014810996,0.0017956856,0.051784236,0.09106804,0.035902653,0.004026952,0.6242226],"study_design_scores_gemma":[0.0001884909,0.0013007745,0.23277652,0.00014466522,0.001019342,0.00094278937,0.000967098,0.6268443,0.043041192,0.08092669,0.011503553,0.00034466668],"about_ca_topic_score_codex":0.0023957882,"about_ca_topic_score_gemma":0.0020034362,"teacher_disagreement_score":0.004841893,"about_ca_system_score_codex":0.00046802528,"about_ca_system_score_gemma":0.0008857732,"threshold_uncertainty_score":0.021998584},"labels":[],"label_agreement":null},{"id":"W4247858078","doi":"10.1109/msr.2015.37","title":"A Method to Detect License Inconsistencies in Large-Scale Open Source Projects","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"License; Header; Computer science; MIT License; Reuse; Software; Source code; Computer security; Resource (disambiguation); Toolbox; Code (set theory); Software engineering; World Wide Web; Data science; Programming language; Engineering; Set (abstract data type); Operating system","score_opus":0.06476595131002189,"score_gpt":0.3443215948206281,"score_spread":0.2795556435106062,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247858078","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04368282,0.00036239927,0.9334248,0.000484969,0.0001796531,0.0011401584,0.0021135246,0.0160459,0.0025658114],"genre_scores_gemma":[0.103830285,0.00013860193,0.8895422,0.00008928788,0.000050586837,0.0005407644,0.0029986908,0.0008277613,0.001981837],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98468876,0.0024454335,0.0025974365,0.0033081665,0.006457656,0.00050245394],"domain_scores_gemma":[0.94695026,0.018336793,0.01134004,0.009342231,0.013125672,0.0009050587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009363319,0.0011196368,0.00094366464,0.012986481,0.0020802806,0.0028585303,0.0026735554,0.0021713204,0.0025332528],"category_scores_gemma":[0.05305306,0.0011635316,0.0011453556,0.007497246,0.0010781222,0.0040077036,0.0036707122,0.0019172637,0.0014358534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006583453,0.00051291205,0.0893121,0.0017541364,0.00033643577,0.0035016797,0.006486107,0.007593354,0.034181844,0.02538583,0.020843327,0.80943394],"study_design_scores_gemma":[0.0005056366,0.0007418908,0.06708968,0.001107271,0.000864017,0.011069737,0.0052315528,0.51194865,0.14165778,0.054457925,0.20442334,0.0009025844],"about_ca_topic_score_codex":0.004419488,"about_ca_topic_score_gemma":0.006623222,"teacher_disagreement_score":0.012986481,"about_ca_system_score_codex":0.0010175718,"about_ca_system_score_gemma":0.004414001,"threshold_uncertainty_score":0.049518585},"labels":[],"label_agreement":null},{"id":"W4248094163","doi":"10.1145/508387.508389","title":"Explicit programming","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Modular programming; Separation of concerns; Programming language; Source code; Vocabulary; Reusability; Code (set theory); Java; Point (geometry); Software engineering; Software; Set (abstract data type)","score_opus":0.03540660668555613,"score_gpt":0.2557462605673754,"score_spread":0.2203396538818193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4248094163","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016072,0.0005486016,0.9565799,0.0011599441,0.0002519651,0.00015130462,0.00025037891,0.0036035948,0.03584704],"genre_scores_gemma":[0.060498558,0.0020560357,0.87678665,0.0016779916,0.00039470615,0.0007551459,0.0010849917,0.0030499962,0.053695846],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99520665,0.00146383,0.0004459385,0.0010662592,0.0014275544,0.00038971542],"domain_scores_gemma":[0.98976445,0.0046637636,0.0005658365,0.003507308,0.0012290165,0.0002696353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005225757,0.001470135,0.0007679514,0.0011584802,0.0016444394,0.005897347,0.0040389025,0.0020102446,0.020740785],"category_scores_gemma":[0.015975837,0.0014300634,0.0019626317,0.0012959451,0.0041964552,0.010847245,0.006606997,0.0050905244,0.0084410645],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005625246,0.00005017326,0.0005069669,0.00047857294,0.00003761504,0.00017275069,0.0010096607,0.0020489725,0.0020442668,0.86921453,0.014746524,0.10963367],"study_design_scores_gemma":[0.00003911786,0.0000443151,0.0001582324,0.00031324296,0.00005989866,0.00053388363,0.00015886444,0.008597838,0.004484062,0.4368342,0.54872674,0.000049735885],"about_ca_topic_score_codex":0.0013157881,"about_ca_topic_score_gemma":0.0016804257,"teacher_disagreement_score":0.020740785,"about_ca_system_score_codex":0.0012293028,"about_ca_system_score_gemma":0.0028195318,"threshold_uncertainty_score":0.06938481},"labels":[],"label_agreement":null},{"id":"W4249156834","doi":"10.22215/etd/2007-07474","title":"Multi-objective genetic algorithm to support class responsibility assignment","year":2007,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Class (philosophy); Computer science; Genealogy; Combinatorics; Humanities; Operations research; Mathematics; Political science; Artificial intelligence; Art; History","score_opus":0.022965209960821267,"score_gpt":0.3427786959922792,"score_spread":0.319813486031458,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4249156834","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05098024,0.00029152786,0.9415202,0.0002963609,0.000105893276,0.00010274968,0.000078688856,0.00097353035,0.0056508826],"genre_scores_gemma":[0.49770412,0.00015571424,0.49559167,0.00012134014,0.00003596518,0.00022289682,0.00018966106,0.00015384641,0.005824735],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996816,0.00012148052,0.000012304426,0.00005124591,0.00008609725,0.000047379308],"domain_scores_gemma":[0.99932146,0.00039766956,0.00006067864,0.000041264168,0.00013362197,0.00004528159],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00090658036,0.00058635196,0.0005674184,0.0008325462,0.0004246718,0.0007601438,0.001183742,0.0008550211,0.0025966219],"category_scores_gemma":[0.0020903163,0.00031589344,0.00043627076,0.0005920003,0.00031230302,0.00055155606,0.00050111837,0.00090097106,0.00036535884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041075833,0.00007402589,0.00041634802,0.000030159184,0.000036074067,0.00004450918,0.000044949225,0.9286567,0.0013888886,0.0038041016,0.00145414,0.06400913],"study_design_scores_gemma":[0.000013179238,0.000013317332,0.000058530473,0.0000036039828,0.0000047945505,0.0000052285236,0.000004382164,0.9986059,0.00024317236,0.0006718294,0.00037435823,0.0000018024391],"about_ca_topic_score_codex":0.010222216,"about_ca_topic_score_gemma":0.009159049,"teacher_disagreement_score":0.010222216,"about_ca_system_score_codex":0.0008576322,"about_ca_system_score_gemma":0.0013710147,"threshold_uncertainty_score":0.020325482},"labels":[],"label_agreement":null},{"id":"W4249157519","doi":"10.1145/1095430.1081744","title":"Strathcona example recommendation tool","year":2005,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Calgary","funders":"","keywords":"Computer science; Documentation; Application programming interface; Software engineering; Code (set theory); Source code; Software; Programming language; Set (abstract data type)","score_opus":0.029946252676475032,"score_gpt":0.2610645669782109,"score_spread":0.23111831430173588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4249157519","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023544142,0.0023841318,0.3904216,0.0043869824,0.0006003461,0.0020341428,0.11761065,0.30243376,0.15658417],"genre_scores_gemma":[0.0752365,0.0020507846,0.67044145,0.0017548903,0.0001844812,0.00217307,0.13160382,0.009612755,0.10694215],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99919766,0.00013360605,0.00008076752,0.00011733474,0.00042487824,0.000045789802],"domain_scores_gemma":[0.9968051,0.0017430596,0.000112940514,0.00036009433,0.0008398739,0.00013900746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00067614723,0.0010425772,0.0007776216,0.0040246374,0.00073354034,0.001338609,0.0014134538,0.0013028578,0.081962496],"category_scores_gemma":[0.0059052547,0.0005493998,0.0006004803,0.003216645,0.000166951,0.0013844142,0.00092553836,0.00071954046,0.02713886],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049522106,0.00014906355,0.0025487302,0.0013101245,0.00006492465,0.00072213117,0.00028832405,0.0025089048,0.0043838187,0.0070622293,0.6263891,0.35407737],"study_design_scores_gemma":[0.00039765143,0.00014922272,0.0038206773,0.00030182852,0.00008084938,0.0009352371,0.00022716347,0.047806773,0.009503877,0.0067755324,0.9298877,0.00011348715],"about_ca_topic_score_codex":0.012361286,"about_ca_topic_score_gemma":0.030937301,"teacher_disagreement_score":0.081962496,"about_ca_system_score_codex":0.0005830278,"about_ca_system_score_gemma":0.0009484979,"threshold_uncertainty_score":0.27419186},"labels":[],"label_agreement":null},{"id":"W4249773613","doi":"10.1109/ase.2004.1342757","title":"Refactoring use case models on episodes","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code refactoring; Computer science; Property (philosophy); Programming language; Software engineering; Artificial intelligence; Software","score_opus":0.08100981101089157,"score_gpt":0.2915194850170447,"score_spread":0.21050967400615317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4249773613","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028013306,0.00023883901,0.9641931,0.0002189719,0.000045056502,0.0003864909,0.0004705368,0.0032954302,0.0031383082],"genre_scores_gemma":[0.22951409,0.0006089002,0.7612282,0.00013142684,0.000048562455,0.00058936345,0.0025961476,0.0011166803,0.004166598],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99482167,0.0017365579,0.00067403406,0.0006214398,0.001802671,0.00034352153],"domain_scores_gemma":[0.98286337,0.007428836,0.0015902421,0.0055548404,0.0022368487,0.00032587323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004778788,0.0012252926,0.00088459626,0.0025659672,0.0005515655,0.002961222,0.0022052096,0.0013916576,0.0022336035],"category_scores_gemma":[0.01883549,0.0011700384,0.0025228194,0.0013125244,0.0011825884,0.0037360305,0.0023714143,0.0021599436,0.0008175319],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004227695,0.00046774588,0.010758009,0.0007721492,0.00039283864,0.002824694,0.003945936,0.4793927,0.03375234,0.22873035,0.007778847,0.23076166],"study_design_scores_gemma":[0.00010116857,0.00019917387,0.0011831927,0.0003352229,0.0002052797,0.000451164,0.00030385063,0.8058359,0.030451378,0.10650679,0.05432165,0.00010525697],"about_ca_topic_score_codex":0.0061326865,"about_ca_topic_score_gemma":0.006232617,"teacher_disagreement_score":0.0061326865,"about_ca_system_score_codex":0.0014661708,"about_ca_system_score_gemma":0.0019827508,"threshold_uncertainty_score":0.025272965},"labels":[],"label_agreement":null},{"id":"W4249834227","doi":"10.1109/se-hpccse.2016.010","title":"Advantages, Disadvantages and Misunderstandings About Document Driven Design for Scientific Software","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada; McMaster University","funders":"","keywords":"Redevelopment; Documentation; Software engineering; Process (computing); Computer science; Software; Software development process; Code (set theory); Process management; Software design; Software development; Engineering management; Engineering; Political science; Programming language; Law","score_opus":0.030273725528896973,"score_gpt":0.2901579016403139,"score_spread":0.25988417611141695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4249834227","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9057321,0.004344038,0.060640484,0.01442891,0.00029230636,0.00041037062,0.000081387894,0.00007397235,0.0139964465],"genre_scores_gemma":[0.9706188,0.001971354,0.024215002,0.0014831293,0.00007898286,0.00039675573,0.000046589088,0.000083156905,0.0011061564],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.75793403,0.18441637,0.011380757,0.004120874,0.039925728,0.0022222162],"domain_scores_gemma":[0.48542595,0.44626513,0.028633367,0.012942523,0.024081646,0.0026513769],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.13785872,0.00064543576,0.0006347443,0.003884414,0.004606495,0.007832246,0.0017055755,0.0020337552,0.00094292726],"category_scores_gemma":[0.25014335,0.0008958312,0.0006066045,0.0029033304,0.012978869,0.01078961,0.0043717865,0.0035170205,0.00021169311],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001296749,0.00010952218,0.019791178,0.0015676388,0.00004543924,0.00081065326,0.8843804,0.00039524405,0.0039477167,0.0204698,0.0007232299,0.06762959],"study_design_scores_gemma":[0.000044235916,0.00048888713,0.0118534425,0.0037188206,0.000074680895,0.0030129128,0.9000304,0.0019104687,0.005429243,0.02105897,0.05222504,0.0001528966],"about_ca_topic_score_codex":0.0010828453,"about_ca_topic_score_gemma":0.0016660708,"teacher_disagreement_score":0.86214125,"about_ca_system_score_codex":0.0061732116,"about_ca_system_score_gemma":0.006738993,"threshold_uncertainty_score":0.7290753},"labels":[],"label_agreement":null},{"id":"W4250021107","doi":"10.1109/icse.2013.6606627","title":"Reverb: Recommending code-related web pages","year":2013,"lang":"en","type":"article","venue":"2013 35th International Conference on Software Engineering (ICSE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"World Wide Web; Computer science; Documentation; Web page; Code (set theory); Source code; Field (mathematics); Web development; Programming language; Set (abstract data type)","score_opus":0.038701204816133136,"score_gpt":0.2726261979321894,"score_spread":0.23392499311605625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250021107","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5093654,0.00656968,0.33295465,0.0017860447,0.0005646099,0.0016134107,0.006072708,0.122461654,0.018611858],"genre_scores_gemma":[0.4520345,0.0014668903,0.51740676,0.0003844171,0.00014543864,0.000330244,0.007150331,0.0019412737,0.019140169],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991992,0.00021702782,0.0000459676,0.00015920596,0.00032989247,0.00004865761],"domain_scores_gemma":[0.99229336,0.004159064,0.00067655253,0.0009887003,0.0014171315,0.00046523596],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013483244,0.0013424059,0.0006854894,0.003540888,0.00052501855,0.0011366875,0.001210928,0.0012861057,0.004850558],"category_scores_gemma":[0.012165769,0.0006689831,0.00038138917,0.0014474309,0.00022022493,0.0015907071,0.00066578126,0.0008814759,0.0038412535],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007941827,0.0015542094,0.038902916,0.0015936327,0.00022080172,0.0010404726,0.0011406647,0.012761086,0.029173015,0.0012188024,0.08594324,0.8256569],"study_design_scores_gemma":[0.00063856115,0.0024933096,0.06375446,0.0004581009,0.0005908331,0.0029605704,0.001329623,0.6539899,0.07486459,0.004709811,0.19381802,0.0003921945],"about_ca_topic_score_codex":0.006977982,"about_ca_topic_score_gemma":0.022083256,"teacher_disagreement_score":0.006977982,"about_ca_system_score_codex":0.00032905734,"about_ca_system_score_gemma":0.0008909921,"threshold_uncertainty_score":0.016226768},"labels":[],"label_agreement":null},{"id":"W4250023757","doi":"10.1109/icse.2015.91","title":"Revisiting the Impact of Classification Techniques on the Performance of Defect Prediction Models","year":2015,"lang":"en","type":"article","venue":"2015 IEEE/ACM 37th IEEE International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":319,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Replicate; Machine learning; Predictive modelling; Artificial intelligence; Software bug; Software; Data mining; Multivariate adaptive regression splines; Multivariate statistics; Set (abstract data type); Variety (cybernetics); Support vector machine; Regression; Logistic regression; Regression analysis; Statistics; Mathematics","score_opus":0.11980003918275281,"score_gpt":0.34096804315866175,"score_spread":0.22116800397590894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250023757","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92608166,0.011943015,0.043994147,0.003549683,0.0009555618,0.00025307943,0.0054328665,0.0028615657,0.0049283523],"genre_scores_gemma":[0.94401,0.0012257398,0.042592756,0.0005736468,0.0003772958,0.00011768015,0.00968091,0.0004227542,0.0009991587],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97400993,0.0121605415,0.0018072542,0.0054308083,0.005303357,0.0012880793],"domain_scores_gemma":[0.81073135,0.13842453,0.007688556,0.026993351,0.014366012,0.0017962275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03935017,0.0029058536,0.0021761896,0.004008595,0.0014939039,0.0048962645,0.0034043833,0.0031377457,0.0011388845],"category_scores_gemma":[0.12628403,0.0007658417,0.003036417,0.0037108562,0.0021200357,0.006605088,0.0028886828,0.006168012,0.0017079423],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005986311,0.002294016,0.34168047,0.0016075384,0.002409377,0.00057871745,0.0013767326,0.21136709,0.012569593,0.002705522,0.030160204,0.38726446],"study_design_scores_gemma":[0.00042637615,0.0022287027,0.10294076,0.0004805963,0.0007865998,0.00051228073,0.0013321079,0.85235,0.019084727,0.008189062,0.0114004705,0.00026823257],"about_ca_topic_score_codex":0.012352693,"about_ca_topic_score_gemma":0.009795201,"teacher_disagreement_score":0.03935017,"about_ca_system_score_codex":0.0016155067,"about_ca_system_score_gemma":0.0019865544,"threshold_uncertainty_score":0.20810604},"labels":[],"label_agreement":null},{"id":"W4250365463","doi":"10.4018/978-1-4666-0261-8.ch020","title":"A Theory of Program Comprehension","year":2012,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"","keywords":"Comprehension; Program comprehension; Computer science; Software; Vision science; Process (computing); Perspective (graphical); Cognitive science; Artificial intelligence; Human–computer interaction; Software system; Psychology; Programming language","score_opus":0.03181813693292335,"score_gpt":0.2792692428518183,"score_spread":0.24745110591889496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250365463","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008130831,0.0076509234,0.6897478,0.013169091,0.00044594932,0.0002578322,0.0006938047,0.0015718958,0.27833188],"genre_scores_gemma":[0.50301206,0.011840759,0.36299658,0.0072118314,0.0018040276,0.001721842,0.0027106416,0.0015122177,0.10719007],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9975757,0.0010342419,0.0001516894,0.00053719536,0.00050943374,0.00019170398],"domain_scores_gemma":[0.994663,0.0037900617,0.00023510963,0.00051070284,0.00066673703,0.00013432282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027078704,0.0012833145,0.000724794,0.0024491944,0.0018813799,0.004938583,0.0022515452,0.0033108187,0.021461792],"category_scores_gemma":[0.008284371,0.0006148155,0.0022199347,0.0019785827,0.008914109,0.015827443,0.0026220332,0.0046766144,0.005134659],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000010206361,0.000021693211,0.00012879755,0.00017116852,0.000009483438,0.00006276906,0.0013498524,0.00085359975,0.0002793472,0.9699943,0.0050164172,0.02210228],"study_design_scores_gemma":[0.000012758118,0.000017758086,0.000138636,0.00011320051,0.000010230516,0.00010918963,0.00021759125,0.0032728037,0.0003461713,0.9473913,0.048359826,0.000010642893],"about_ca_topic_score_codex":0.002648078,"about_ca_topic_score_gemma":0.0011489284,"teacher_disagreement_score":0.021461792,"about_ca_system_score_codex":0.0028924926,"about_ca_system_score_gemma":0.0023937821,"threshold_uncertainty_score":0.071796834},"labels":[],"label_agreement":null},{"id":"W4250422959","doi":"10.1145/633482.633495","title":"\"Bloat\"","year":2000,"lang":"en","type":"article","venue":"CHI '00 extended abstracts on Human factors in computer systems - CHI '00","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Word (group theory); Software; Term (time); World Wide Web; Programming language; Linguistics","score_opus":0.03712781445241532,"score_gpt":0.2910802495229868,"score_spread":0.25395243507057147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250422959","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11504093,0.012711181,0.08296361,0.027597956,0.006498071,0.00047186273,0.0022076978,0.003853427,0.7486553],"genre_scores_gemma":[0.6665303,0.005506812,0.025082985,0.03520164,0.0016410877,0.0005325134,0.0022085502,0.00206723,0.26122883],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.994093,0.0024031664,0.00033631263,0.0006845334,0.0017888049,0.00069415674],"domain_scores_gemma":[0.9924257,0.0016908336,0.001923059,0.0015187744,0.00188569,0.0005560033],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031072088,0.00090317783,0.00042678518,0.0023899204,0.0047867117,0.0056520645,0.0013984964,0.0025096722,0.023605447],"category_scores_gemma":[0.012772126,0.0003623767,0.000659045,0.0026919881,0.007871451,0.008766645,0.006770213,0.002761676,0.010963015],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024826845,0.00007563103,0.008338573,0.00090696017,0.00007371366,0.0007493877,0.055294894,0.0002089473,0.0024889074,0.41731519,0.3710664,0.14323309],"study_design_scores_gemma":[0.000010473549,0.000057217538,0.0052028275,0.00033490194,0.000026830188,0.0015189796,0.010818776,0.00029318087,0.0006614164,0.016742766,0.96427864,0.000054132815],"about_ca_topic_score_codex":0.007834681,"about_ca_topic_score_gemma":0.011443978,"teacher_disagreement_score":0.023605447,"about_ca_system_score_codex":0.002168901,"about_ca_system_score_gemma":0.0018439972,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4250612015","doi":"10.4018/978-1-60566-170-4.ch020","title":"Constructivist Learning During Software Development","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Laurentian University","funders":"","keywords":"Constructivist teaching methods; Programmer; Denial; Computer science; Process (computing); Documentation; Mathematics education; Cognitive science; Knowledge management; Psychology; Epistemology; Programming language; Philosophy; Teaching method","score_opus":0.018925662165124438,"score_gpt":0.22807582701891668,"score_spread":0.20915016485379223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250612015","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03133198,0.04511382,0.36255112,0.0165241,0.00049905275,0.00012020842,0.00006434054,0.0003319029,0.5434634],"genre_scores_gemma":[0.6392998,0.039397985,0.15729253,0.0021445616,0.0004941037,0.00037447808,0.00016229125,0.00028788563,0.16054645],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9986702,0.00074764766,0.000035414858,0.00015164104,0.00032747688,0.00006758925],"domain_scores_gemma":[0.99767953,0.0019565793,0.00006732865,0.00016401988,0.000085552296,0.000047086152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021800294,0.000590943,0.00027091507,0.0007944271,0.00092334946,0.0030666867,0.0013161359,0.001534007,0.0030542442],"category_scores_gemma":[0.0029452762,0.00035784443,0.00030205425,0.0010293454,0.011995663,0.0051268903,0.0021970472,0.0039552683,0.00074980967],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000132922205,0.000029723185,0.00023923523,0.00026185345,0.0000070661677,0.00011558283,0.012800407,0.00258904,0.00040043888,0.9061869,0.004133965,0.073222585],"study_design_scores_gemma":[0.000017843127,0.000035937737,0.00033288033,0.00042507396,0.000006510597,0.0002106865,0.0012427588,0.0031383606,0.0008298019,0.7955281,0.19821754,0.000014422821],"about_ca_topic_score_codex":0.0016609004,"about_ca_topic_score_gemma":0.0023793536,"teacher_disagreement_score":0.00417089,"about_ca_system_score_codex":0.00417089,"about_ca_system_score_gemma":0.0020561893,"threshold_uncertainty_score":0.030262113},"labels":[],"label_agreement":null},{"id":"W4250612791","doi":"10.7287/peerj.preprints.1920","title":"Analyzing test driven development based on GitHub evidence","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Commit; Agile software development; Computer science; Test-driven development; Software engineering; Java; Process (computing); Software; Software development; Database; Programming language","score_opus":0.04884592265832392,"score_gpt":0.3013920882778682,"score_spread":0.25254616561954424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250612791","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9728247,0.0024451104,0.01197372,0.0008028942,0.000030459038,0.00033202255,0.004415322,0.00025908338,0.0069166822],"genre_scores_gemma":[0.9791007,0.0009556451,0.010695545,0.00020217286,0.0000362128,0.00038800974,0.007682703,0.00017934991,0.00075955765],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9540083,0.015827417,0.004346623,0.003401757,0.02086429,0.0015516235],"domain_scores_gemma":[0.5113423,0.33389348,0.058702853,0.04670642,0.046914473,0.0024404787],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.023713054,0.00065040606,0.00076169334,0.024400178,0.0008035073,0.0039068027,0.0024964735,0.0012022754,0.0018610902],"category_scores_gemma":[0.26709527,0.000549934,0.0010033724,0.027299372,0.002271371,0.0030986539,0.0033501608,0.0012089123,0.00049928133],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013013989,0.0005639406,0.81315607,0.002809653,0.0010652835,0.0020603554,0.0053548203,0.007759248,0.0040472047,0.009764862,0.005545318,0.14657192],"study_design_scores_gemma":[0.00021359966,0.0010214731,0.9268218,0.0017637897,0.00079162954,0.0019954112,0.0054873135,0.02428982,0.008819428,0.0064734677,0.02217319,0.00014912657],"about_ca_topic_score_codex":0.009155216,"about_ca_topic_score_gemma":0.00953601,"teacher_disagreement_score":0.97628695,"about_ca_system_score_codex":0.0021316765,"about_ca_system_score_gemma":0.0023649784,"threshold_uncertainty_score":0.12540811},"labels":[],"label_agreement":null},{"id":"W4250654813","doi":"10.1007/s10664-007-9041-9","title":"In this issue","year":2007,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science","score_opus":0.018849793462910394,"score_gpt":0.30064744145276495,"score_spread":0.28179764798985457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250654813","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0007511217,0.009938868,0.0012786861,0.22083767,0.58960927,0.000089364075,0.00048288712,0.000465004,0.1765471],"genre_scores_gemma":[0.0039857044,0.00520246,0.00061968266,0.09185476,0.21899368,0.000070594484,0.00057668635,0.0003167605,0.67837965],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983438,0.0002290861,0.0000940467,0.00023383199,0.0008493219,0.00025001125],"domain_scores_gemma":[0.99286026,0.0017164533,0.0003938139,0.00080279866,0.002136507,0.00209021],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.002574584,0.0008418999,0.000912213,0.0022071104,0.0033259508,0.009660361,0.0021291634,0.0074923863,0.21002243],"category_scores_gemma":[0.011808608,0.00047851843,0.0009396828,0.0014096851,0.0017252048,0.0038372288,0.002521483,0.0075734593,0.10152897],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000014415005,0.000035307017,0.00023460817,0.00004798298,0.0000036725064,0.000031026768,0.000028428512,0.000008417489,0.00008082844,0.0018891385,0.9798364,0.017789746],"study_design_scores_gemma":[0.000007983311,0.00000883896,0.0004318654,0.000058596306,0.0000053522017,0.00003732751,0.00008188314,0.00002585187,0.000081058555,0.0013601817,0.99789643,0.0000045912857],"about_ca_topic_score_codex":0.001977279,"about_ca_topic_score_gemma":0.008431475,"teacher_disagreement_score":0.78997755,"about_ca_system_score_codex":0.0014676962,"about_ca_system_score_gemma":0.0033502246,"threshold_uncertainty_score":0.702595},"labels":[],"label_agreement":null},{"id":"W4250972184","doi":"10.1109/icse.2013.6606676","title":"Why did this code change?","year":2013,"lang":"en","type":"article","venue":"2013 35th International Conference on Software Engineering (ICSE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Commit; Automatic summarization; Documentation; Code (set theory); Code review; Source code; Set (abstract data type); World Wide Web; Programming language; Static program analysis; Information retrieval; Database; Software; Software development","score_opus":0.0570751759543973,"score_gpt":0.28208227329182556,"score_spread":0.22500709733742824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250972184","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34580272,0.0066514164,0.15891652,0.2789614,0.016889906,0.0013013749,0.006766381,0.007941656,0.17676874],"genre_scores_gemma":[0.7903099,0.0025429847,0.054949254,0.02340621,0.0011763566,0.0002660619,0.003550626,0.0023034294,0.12149516],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99807334,0.0004658231,0.00012376525,0.00044115825,0.0006759525,0.00021998366],"domain_scores_gemma":[0.9923064,0.0019532302,0.0012890704,0.00065165403,0.0031902515,0.00060948456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020288697,0.00048398308,0.00032024118,0.00107997,0.0020923882,0.001947669,0.0005554847,0.0019127332,0.009195253],"category_scores_gemma":[0.018071417,0.0003365597,0.0006054057,0.00082908024,0.0016045207,0.003250093,0.0010634491,0.002579061,0.0033915655],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005617839,0.0002446398,0.08765548,0.0012745984,0.00018788676,0.013435229,0.042035986,0.0011207495,0.026443169,0.07162537,0.3182879,0.43712723],"study_design_scores_gemma":[0.00005932461,0.0002804985,0.05081233,0.000880711,0.00018307193,0.010860378,0.02503778,0.0049810763,0.014105621,0.030890433,0.861682,0.00022671916],"about_ca_topic_score_codex":0.011166422,"about_ca_topic_score_gemma":0.021038624,"teacher_disagreement_score":0.011166422,"about_ca_system_score_codex":0.0022685444,"about_ca_system_score_gemma":0.0023905274,"threshold_uncertainty_score":0.030761123},"labels":[],"label_agreement":null},{"id":"W4251091893","doi":"10.22215/etd/2011-08906","title":"Maintenance of hybrid software that combines open source and closed source components","year":2011,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Canadian Heritage; Library and Archives Canada","funders":"","keywords":"Open source; Open source software; Computer science; Software; Operating system","score_opus":0.03346099502205735,"score_gpt":0.26784044478024405,"score_spread":0.2343794497581867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4251091893","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47377545,0.0009497798,0.49725297,0.00041289456,0.00012719758,0.00020163092,0.0002532909,0.018935088,0.008091774],"genre_scores_gemma":[0.7699889,0.00024088197,0.21605462,0.00011633548,0.00008273672,0.00016708305,0.0012934519,0.002048487,0.010007588],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973495,0.00047523144,0.00025290746,0.00043664972,0.0013558078,0.00012980303],"domain_scores_gemma":[0.9844297,0.0039199037,0.0013522791,0.0070624514,0.0027694728,0.00046628452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029296374,0.0004434982,0.00036898904,0.0019888624,0.0005654718,0.0024569845,0.0021402368,0.0007492975,0.0014272254],"category_scores_gemma":[0.0109463325,0.00068434776,0.0005065929,0.0014359772,0.0008328616,0.0037312221,0.0019063106,0.00082876114,0.00069583167],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010197273,0.000840169,0.025552075,0.00048260606,0.00037595327,0.0009994444,0.0019953568,0.024759026,0.15586856,0.013084707,0.0072985375,0.7677238],"study_design_scores_gemma":[0.0005537284,0.0023607048,0.04823165,0.0002925726,0.000976099,0.0034050182,0.0009240251,0.63293797,0.21023501,0.029888133,0.07000257,0.0001926121],"about_ca_topic_score_codex":0.0013897232,"about_ca_topic_score_gemma":0.0017236482,"teacher_disagreement_score":0.0029296374,"about_ca_system_score_codex":0.000606562,"about_ca_system_score_gemma":0.00088814547,"threshold_uncertainty_score":0.015493631},"labels":[],"label_agreement":null},{"id":"W4251625061","doi":"10.1109/icse.2015.291","title":"Towards Generation of Software Development Tasks","year":2015,"lang":"en","type":"article","venue":"2015 IEEE/ACM 37th IEEE International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Task (project management); Decomposition; Context (archaeology); Process (computing); Software engineering; Task analysis; Software development; Software; Human–computer interaction; Programming language; Systems engineering; Engineering","score_opus":0.14010299695419845,"score_gpt":0.3348344149472631,"score_spread":0.19473141799306468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4251625061","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012451479,0.00021191993,0.9787998,0.00050663727,0.00011894344,0.00033564502,0.0007923311,0.0029160273,0.0038671833],"genre_scores_gemma":[0.06457659,0.00022845178,0.92772466,0.00012336549,0.000039422765,0.0003283036,0.0034246386,0.0008964556,0.002658102],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99531823,0.0017954509,0.00038077342,0.000988766,0.0012725353,0.00024424194],"domain_scores_gemma":[0.9884183,0.005166571,0.0007490962,0.00253828,0.0027943258,0.0003334203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033736115,0.0014343883,0.0005989593,0.002499218,0.0010285636,0.002719719,0.0017539331,0.0015978097,0.004226527],"category_scores_gemma":[0.022308104,0.0011289648,0.0023572464,0.0014183187,0.0009408762,0.00258005,0.003339974,0.002823261,0.0038965975],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036633003,0.0003676942,0.008715352,0.0015000198,0.00013354137,0.0013986877,0.005278615,0.06744083,0.02402068,0.13550237,0.0286989,0.72657704],"study_design_scores_gemma":[0.00010111402,0.00018700155,0.0025176308,0.00069054624,0.000112937014,0.0007047606,0.0012441883,0.6558204,0.026023922,0.20355494,0.10893388,0.00010878815],"about_ca_topic_score_codex":0.0027024238,"about_ca_topic_score_gemma":0.003288358,"teacher_disagreement_score":0.004226527,"about_ca_system_score_codex":0.0011746024,"about_ca_system_score_gemma":0.0035335217,"threshold_uncertainty_score":0.017841578},"labels":[],"label_agreement":null},{"id":"W4251684089","doi":"10.22215/etd/2021-14511","title":"Recommending GitHub Projects by Leveraging Developers' Social Networks and Genetic Algorithm","year":2021,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Construct (python library); Data science; Genetic algorithm; World Wide Web; Social network (sociolinguistics); Software engineering; Knowledge management; Social media; Machine learning; Programming language","score_opus":0.019088986061369852,"score_gpt":0.2709855241976485,"score_spread":0.25189653813627866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4251684089","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5734434,0.0033850505,0.39634448,0.002527694,0.00048274148,0.00073101465,0.0010936355,0.0045514856,0.01744054],"genre_scores_gemma":[0.805583,0.0009743752,0.18070625,0.0003738547,0.00023686203,0.00029864485,0.001848222,0.0004935739,0.009485303],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99886334,0.00031575357,0.00004931855,0.00032253907,0.00034509136,0.00010397921],"domain_scores_gemma":[0.9965702,0.0017996546,0.00036293018,0.00031487373,0.00069935713,0.00025290175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014807768,0.0009446123,0.00083584775,0.004618705,0.0008873474,0.0014692729,0.0011539513,0.001273403,0.0016247744],"category_scores_gemma":[0.008689653,0.00043466303,0.00076838006,0.003042441,0.00039955578,0.0019557227,0.00092125934,0.0008615325,0.0009767392],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003742785,0.0010987571,0.065595716,0.00046641633,0.00046548934,0.0003494753,0.0006675921,0.15365411,0.008948376,0.0043363916,0.028470526,0.7355728],"study_design_scores_gemma":[0.00004613704,0.00013624493,0.005836072,0.000035207595,0.00012451752,0.000095948904,0.00020719589,0.98320353,0.0019544952,0.0049739894,0.0033589671,0.000027588861],"about_ca_topic_score_codex":0.012500043,"about_ca_topic_score_gemma":0.029735094,"teacher_disagreement_score":0.012500043,"about_ca_system_score_codex":0.0008161437,"about_ca_system_score_gemma":0.0016252749,"threshold_uncertainty_score":0.0248546},"labels":[],"label_agreement":null},{"id":"W4251686833","doi":"10.1109/icre.2004.1335682","title":"Helping analysts trace requirements:an objective look","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University; National Aeronautics and Space Administration","keywords":"Tracing; Computer science; TRACE (psycholinguistics); Requirements analysis; Software engineering; Process (computing); Requirements management; Software requirements specification; Systems engineering; Engineering; Programming language; Software; Software design; Software development","score_opus":0.029755968137820665,"score_gpt":0.306384623336623,"score_spread":0.2766286551988023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4251686833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17529517,0.007953822,0.7585214,0.024652306,0.000526825,0.0010838661,0.00046832394,0.004101689,0.027396645],"genre_scores_gemma":[0.46303272,0.005134863,0.52112734,0.0013349403,0.00035393654,0.00060941663,0.00043529648,0.0006893406,0.0072821663],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.92060643,0.039613765,0.005149171,0.002713069,0.030914538,0.0010030308],"domain_scores_gemma":[0.7666728,0.11545672,0.029867105,0.01802258,0.06678503,0.0031957978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.086146206,0.0014738924,0.0012106022,0.005394603,0.0020529698,0.013873372,0.002579141,0.0023223856,0.0022901283],"category_scores_gemma":[0.1685391,0.0011834403,0.000561516,0.003118262,0.002247996,0.020107398,0.0042685037,0.0028675867,0.0010041404],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041504257,0.0012828998,0.035384607,0.002648132,0.00024097408,0.00025802696,0.01795984,0.0035215253,0.017577669,0.019572277,0.0096069,0.8915322],"study_design_scores_gemma":[0.0012685864,0.012095819,0.10322749,0.013082289,0.001745701,0.0054440508,0.1298049,0.14768392,0.13526405,0.1066284,0.34232605,0.001428698],"about_ca_topic_score_codex":0.00097093993,"about_ca_topic_score_gemma":0.0025130732,"teacher_disagreement_score":0.086146206,"about_ca_system_score_codex":0.0019290579,"about_ca_system_score_gemma":0.0044169836,"threshold_uncertainty_score":0.45559013},"labels":[],"label_agreement":null},{"id":"W4251941390","doi":"10.1177/0265532220929918","title":"Automated scoring of junior and senior high essays using Coh-Metrix features: Implications for large-scale language testing","year":2020,"lang":"en","type":"article","venue":"Language Testing","topic":"Software Engineering Research","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Advancing Health Outcomes; University of Alberta","funders":"","keywords":"Natural language processing; Artificial intelligence; Rating scale; Computer science; Scale (ratio); Quality (philosophy); Construct (python library); Psychology; Computational linguistics; Disadvantaged; Machine learning; Developmental psychology","score_opus":0.0384669575825602,"score_gpt":0.314934657209404,"score_spread":0.2764676996268438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4251941390","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9057022,0.00015473893,0.08502079,0.000737787,0.000078240904,0.0006546102,0.0009435304,0.0016023803,0.005105739],"genre_scores_gemma":[0.9451553,0.000028603572,0.05261295,0.00005831062,0.000029436469,0.00038280082,0.00070253544,0.000082795756,0.000947226],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9826245,0.011149075,0.001145939,0.0013772713,0.003304211,0.00039899288],"domain_scores_gemma":[0.87404466,0.0725709,0.013568377,0.014366395,0.02334865,0.0021009862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022005778,0.000703687,0.0005234812,0.0018903319,0.00061810634,0.0018867017,0.0011379193,0.0005495879,0.0016732344],"category_scores_gemma":[0.1303999,0.00025200358,0.0004401214,0.0019251253,0.00070892076,0.0020835954,0.0017141878,0.0012150668,0.00063632644],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000787052,0.00081563153,0.51557225,0.00023529692,0.0001439789,0.00012137948,0.0030621337,0.010712434,0.00787936,0.002254873,0.006686268,0.4517294],"study_design_scores_gemma":[0.00013234996,0.0017263475,0.665396,0.00016275565,0.00006515012,0.00031676202,0.0032229463,0.29533166,0.018650383,0.0063330657,0.008492175,0.00017050953],"about_ca_topic_score_codex":0.0041532195,"about_ca_topic_score_gemma":0.009832878,"teacher_disagreement_score":0.022005778,"about_ca_system_score_codex":0.0011107483,"about_ca_system_score_gemma":0.001674405,"threshold_uncertainty_score":0.11637908},"labels":[],"label_agreement":null},{"id":"W4252165735","doi":"10.24124/2006/bpgub1329","title":"Distribution of defects in a large software system","year":2006,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Memory leak; Source lines of code; Programmer; Leak; Software; Operating system; Embedded system; Memory management; Engineering; Semiconductor memory","score_opus":0.007315300180494949,"score_gpt":0.251122040824903,"score_spread":0.24380674064440808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252165735","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99825567,0.000088205525,0.0006658855,0.000036505808,0.0000013037584,0.000015767124,0.00027453402,0.00006895597,0.00059305603],"genre_scores_gemma":[0.99874854,0.000055344808,0.0005594473,0.000011476741,0.0000023227328,0.000009719232,0.0002922964,0.000013388068,0.00030736194],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9975823,0.00031009063,0.0002449185,0.000441983,0.0012157013,0.00020486437],"domain_scores_gemma":[0.9598323,0.015545284,0.012310017,0.003012611,0.008390132,0.0009097311],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013002272,0.00015270585,0.00021269683,0.0068292795,0.00059687893,0.0008493128,0.00044124684,0.000389657,0.0009890539],"category_scores_gemma":[0.017194629,0.00025949551,0.00017140507,0.0030065847,0.0010239541,0.00062142184,0.0007442502,0.00041845744,0.00019732665],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016620834,0.000071343646,0.96119285,0.00005253517,0.00003213414,0.0007397294,0.0028831325,0.0012043675,0.004602725,0.0002841704,0.00032860923,0.028442163],"study_design_scores_gemma":[0.00000518957,0.00031168276,0.99056506,0.000026863388,0.000039565974,0.0011751895,0.0018667288,0.0019979766,0.002919161,0.0001717455,0.00089773996,0.00002310404],"about_ca_topic_score_codex":0.020792356,"about_ca_topic_score_gemma":0.021630907,"teacher_disagreement_score":0.020792356,"about_ca_system_score_codex":0.0012001118,"about_ca_system_score_gemma":0.0009317948,"threshold_uncertainty_score":0.041342676},"labels":[],"label_agreement":null},{"id":"W4252238268","doi":"10.4018/9781591409411.ch008.ch000","title":"Modeling Relevance Relations Using Machine Learning Techniques","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Relevance (law); Relation (database); Artificial intelligence; Software deployment; Software; Machine learning; Abstraction; Precision and recall; Tuple; Data mining; Information retrieval; Data science; Software engineering; Programming language","score_opus":0.04127005045233979,"score_gpt":0.27500254848819566,"score_spread":0.23373249803585588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252238268","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024038525,0.0034776195,0.96192294,0.0015093792,0.00010859195,0.00025020883,0.0006701801,0.0012795704,0.0067429077],"genre_scores_gemma":[0.41753078,0.0030247003,0.56881976,0.00051029574,0.00044128063,0.0006333013,0.0023536761,0.0003599689,0.0063262153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937289,0.002356597,0.0004315272,0.0012635421,0.0018900317,0.00032931744],"domain_scores_gemma":[0.9791027,0.017297586,0.0011903353,0.0010273581,0.0012067292,0.00017533886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059495433,0.001417273,0.001514208,0.0067682574,0.0013945411,0.0037175966,0.0026358627,0.001982934,0.004241503],"category_scores_gemma":[0.03148562,0.0009906853,0.0020706465,0.0051889718,0.0013389545,0.0077887233,0.002034723,0.0032247764,0.0016743571],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019689974,0.00043816754,0.0135004325,0.0007046642,0.0003054164,0.0008802291,0.0012525,0.3963636,0.0022949637,0.14696962,0.017218135,0.4198754],"study_design_scores_gemma":[0.000015048234,0.000029172515,0.0011540798,0.000078997575,0.000040183542,0.00018432798,0.0000834562,0.85266596,0.00062804,0.13883781,0.0062589864,0.000023917353],"about_ca_topic_score_codex":0.007233726,"about_ca_topic_score_gemma":0.0077275927,"teacher_disagreement_score":0.007233726,"about_ca_system_score_codex":0.0023671484,"about_ca_system_score_gemma":0.0016945815,"threshold_uncertainty_score":0.031464577},"labels":[],"label_agreement":null},{"id":"W4252288227","doi":"10.7287/peerj.preprints.3123","title":"Finding and correcting syntax errors using recurrent neural networks","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Syntax error; Computer science; Syntax; Security token; Abstract syntax; Programming language; Parsing; Artificial intelligence; Natural language processing; Language model; Abstract syntax tree","score_opus":0.0793394262160843,"score_gpt":0.3405356839980847,"score_spread":0.2611962577820004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252288227","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28565705,0.0009843677,0.66222,0.0012343476,0.0003413619,0.00011390581,0.0010040441,0.045342036,0.0031027973],"genre_scores_gemma":[0.7287058,0.00031666318,0.2639612,0.00032659507,0.000050549574,0.00008833617,0.0017832684,0.001033983,0.0037335465],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984377,0.00041037233,0.00010888867,0.0005677611,0.00031665404,0.00015866736],"domain_scores_gemma":[0.9945392,0.002575825,0.0008015835,0.00072146195,0.0012155994,0.00014641225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015551428,0.0018917694,0.0006824512,0.001571362,0.0004402295,0.0011870718,0.0020242252,0.0012218545,0.0016221313],"category_scores_gemma":[0.013519768,0.00079236814,0.000812367,0.0008925853,0.0005842352,0.002465445,0.0011446221,0.0017231962,0.0011771859],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005394867,0.00029008152,0.017241376,0.00045359312,0.00029696457,0.0012392715,0.001023805,0.1985671,0.06001181,0.003257736,0.014075679,0.70300305],"study_design_scores_gemma":[0.000019676412,0.000066121,0.0015795374,0.000043372143,0.00006521432,0.000103743056,0.00010972059,0.9766481,0.015805982,0.003842176,0.0016866178,0.000029744577],"about_ca_topic_score_codex":0.011218475,"about_ca_topic_score_gemma":0.0169805,"teacher_disagreement_score":0.011218475,"about_ca_system_score_codex":0.0011742981,"about_ca_system_score_gemma":0.001447588,"threshold_uncertainty_score":0.022306383},"labels":[],"label_agreement":null},{"id":"W4252721210","doi":"10.1109/raise.2012.6227963","title":"Clone detection meets Semantic Web-based transitive closure computation","year":2012,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Source code; Transitive relation; Transitive closure; Semantic Web; Social Semantic Web; Programming language; clone (Java method); Theoretical computer science; Web service; Data mining; Artificial intelligence; Mathematics","score_opus":0.014909508081670238,"score_gpt":0.25853503482566575,"score_spread":0.2436255267439955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252721210","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04853594,0.00011599832,0.9443918,0.00042987213,0.00004261093,0.00022400409,0.00020504904,0.0020754456,0.0039793695],"genre_scores_gemma":[0.5482216,0.00016275783,0.4481102,0.00019718336,0.00009648684,0.00031031106,0.0008451761,0.00036864603,0.0016877032],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98963296,0.0022024887,0.00084072823,0.002619159,0.004055237,0.0006495003],"domain_scores_gemma":[0.96702284,0.020884752,0.0027891065,0.0048354063,0.003952295,0.0005155064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005155579,0.0008336508,0.0015274696,0.004669369,0.0018875342,0.004421458,0.0018256215,0.0014561843,0.0022065928],"category_scores_gemma":[0.03786172,0.0006605135,0.0031133266,0.00208422,0.004013483,0.00847498,0.0035148922,0.0019325734,0.00057549635],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006403879,0.00052133773,0.012227818,0.00073860295,0.00029205796,0.0016860217,0.0025624363,0.08576125,0.0279785,0.53574,0.0045990963,0.32725245],"study_design_scores_gemma":[0.000066104214,0.00012035287,0.0012860218,0.00008274438,0.00013350415,0.0005291422,0.00040251526,0.50210184,0.026946494,0.46276346,0.0055018547,0.0000660033],"about_ca_topic_score_codex":0.005114632,"about_ca_topic_score_gemma":0.0040157535,"teacher_disagreement_score":0.005155579,"about_ca_system_score_codex":0.0022866118,"about_ca_system_score_gemma":0.0028059953,"threshold_uncertainty_score":0.027265608},"labels":[],"label_agreement":null},{"id":"W4252745723","doi":"10.32920/ryerson.14668263.v1","title":"Researchgate.net crawler and a new contribution determines sequence (CDS) method","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Web crawler; Crawling; Scripting language; Computer science; Focused crawler; World Wide Web; Field (mathematics); Java; Information retrieval; Data mining; Data science; The Internet; Web server; Static web page; Operating system; Mathematics","score_opus":0.09491528026853427,"score_gpt":0.3891979685653524,"score_spread":0.29428268829681814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252745723","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022547405,0.0011166211,0.85854876,0.000742081,0.0005572446,0.001701369,0.017507792,0.08375985,0.013518895],"genre_scores_gemma":[0.04776071,0.00040012304,0.91548795,0.00012169325,0.00013605584,0.0010617373,0.017169088,0.005586522,0.012276128],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952154,0.00093351625,0.0007096022,0.0012035302,0.0017681387,0.00016974543],"domain_scores_gemma":[0.9920763,0.0032285678,0.00046814742,0.0015209409,0.00237523,0.00033086457],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038544703,0.0015301746,0.001320596,0.014347897,0.0010181516,0.0030215497,0.0014272382,0.0014041221,0.008224028],"category_scores_gemma":[0.02053499,0.0010242581,0.0013362896,0.008686018,0.0006155015,0.0032483123,0.001927129,0.0011354802,0.0055214465],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005574642,0.00031134882,0.013140412,0.00177803,0.00040752022,0.00043869202,0.0008898772,0.009458526,0.009064628,0.031068306,0.10530045,0.8275847],"study_design_scores_gemma":[0.00072005077,0.00034224562,0.014109539,0.00032689652,0.00039741912,0.0016809563,0.000490335,0.41936758,0.03220912,0.03923316,0.49084386,0.0002788002],"about_ca_topic_score_codex":0.008970808,"about_ca_topic_score_gemma":0.013183912,"teacher_disagreement_score":0.014347897,"about_ca_system_score_codex":0.001391957,"about_ca_system_score_gemma":0.0040285443,"threshold_uncertainty_score":0.027512133},"labels":[],"label_agreement":null},{"id":"W4252859673","doi":"10.4236/ib.2012.51a006","title":"Forecasting and the Role of Churn in Software-as-a-Service Business Models","year":2013,"lang":"en","type":"article","venue":"iBusiness","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software as a service; Revenue; Business model; Computer science; Service (business); Software; Key (lock); Revenue model; Process management; Business; Knowledge management; Marketing; Software development; Finance; Computer security","score_opus":0.01880348607987751,"score_gpt":0.21835534439111914,"score_spread":0.19955185831124164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252859673","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45905942,0.0021161574,0.51644754,0.0050162515,0.00015809818,0.00008178484,0.00028109344,0.000498622,0.016341079],"genre_scores_gemma":[0.9856909,0.00061294425,0.011811695,0.00007922997,0.00006819663,0.000028141707,0.00009428005,0.000038114493,0.0015765283],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99872285,0.0006493785,0.000052560717,0.00017059274,0.00023606521,0.0001685064],"domain_scores_gemma":[0.98808783,0.009520933,0.0009995839,0.00029196913,0.0007746293,0.00032507055],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004195194,0.00084629754,0.0008685684,0.0014277425,0.0009209574,0.0022807275,0.0013016125,0.0013590391,0.0011135129],"category_scores_gemma":[0.01942655,0.0005864285,0.000539073,0.0014275218,0.0010607137,0.0029239934,0.0010128702,0.0019158262,0.00019881748],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003722505,0.000037935904,0.007144017,0.000024445935,0.000028642953,0.00009947291,0.00020008929,0.947338,0.00019760999,0.0331588,0.0006969723,0.011036848],"study_design_scores_gemma":[0.0000012000214,0.000008018975,0.0003570894,0.0000073088636,0.0000029141381,0.000013303926,0.000023679182,0.99013174,0.000034777873,0.009215048,0.00019938969,0.000005483564],"about_ca_topic_score_codex":0.025302231,"about_ca_topic_score_gemma":0.016349535,"teacher_disagreement_score":0.025302231,"about_ca_system_score_codex":0.002257536,"about_ca_system_score_gemma":0.0012965068,"threshold_uncertainty_score":0.050309896},"labels":[],"label_agreement":null},{"id":"W4252902208","doi":"10.1145/566180.566183","title":"Investigating the use of analysis contracts to support fault isolation in object oriented code","year":2002,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Isolation (microbiology); Precondition; Design by contract; Software engineering; Reuse; Object-oriented programming; Software; Object (grammar); Code reuse; Code (set theory); Fault detection and isolation; Instrumentation (computer programming); Programming language; Reliability engineering; Software system; Software construction; Engineering; Set (abstract data type); Artificial intelligence","score_opus":0.054195060972924855,"score_gpt":0.2721160425029615,"score_spread":0.21792098153003664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252902208","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50128204,0.00048183935,0.49070108,0.0012133193,0.000029083278,0.00035594503,0.00004329649,0.0010606002,0.00483281],"genre_scores_gemma":[0.7894636,0.00031669677,0.20894764,0.00013479253,0.00001628631,0.0001460364,0.00005184641,0.00013948811,0.0007836373],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98974216,0.0065993345,0.00037680467,0.0004197521,0.002288448,0.000573453],"domain_scores_gemma":[0.90041894,0.0791427,0.007714095,0.007379305,0.004720452,0.0006245914],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014083331,0.000560356,0.00032749158,0.0008148984,0.0007765619,0.0012669137,0.0015525956,0.0013397417,0.00093993224],"category_scores_gemma":[0.05682043,0.0004858882,0.00035293086,0.0009902541,0.0018239425,0.0045235828,0.001390593,0.001436192,0.00014182551],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012397077,0.0024398386,0.0824853,0.0012789048,0.0002104694,0.0016909628,0.012509607,0.21907403,0.09742058,0.15310591,0.0017591962,0.42678553],"study_design_scores_gemma":[0.00022109045,0.0020807663,0.014329245,0.0003273431,0.0002272063,0.0008671781,0.002857544,0.7996668,0.12661235,0.03616878,0.016523844,0.000117938405],"about_ca_topic_score_codex":0.002270296,"about_ca_topic_score_gemma":0.0017799109,"teacher_disagreement_score":0.014083331,"about_ca_system_score_codex":0.00085569354,"about_ca_system_score_gemma":0.0021924295,"threshold_uncertainty_score":0.07448065},"labels":[],"label_agreement":null},{"id":"W4252955832","doi":"10.5194/gmdd-5-347-2012","title":"Assessing climate model software quality: a defect density analysis of three models","year":2012,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Executable; Context (archaeology); Software; Climate model; Computer science; Trustworthiness; Software quality; Quality (philosophy); Climate change; Software development; Geography; Geology; Programming language; Computer security","score_opus":0.1484321936416739,"score_gpt":0.37334999201768193,"score_spread":0.22491779837600803,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252955832","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9933569,0.000052393265,0.0058187465,0.00008161915,0.0000031854265,0.000020936279,0.00012959173,0.00008941476,0.0004472009],"genre_scores_gemma":[0.99756193,0.000022063372,0.0021307655,0.0000069152047,0.0000024848105,0.000014628715,0.00019389742,0.000015657752,0.00005173698],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9977738,0.0011479448,0.00014391608,0.00027950926,0.0005071517,0.00014758785],"domain_scores_gemma":[0.92149603,0.06264917,0.0050279833,0.00523723,0.004447084,0.0011424542],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008777562,0.00066100695,0.0006054104,0.002576587,0.0005241722,0.0016316792,0.001191479,0.0008782675,0.000728225],"category_scores_gemma":[0.034270853,0.00040598205,0.0017047437,0.0017919154,0.0011668836,0.0015899166,0.0011029515,0.0010107057,0.000098155455],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066719763,0.0006365416,0.46737078,0.000093179035,0.000528842,0.00018734446,0.00067064655,0.5094769,0.0018504277,0.004161351,0.0008800733,0.01347668],"study_design_scores_gemma":[0.000026136988,0.00019003742,0.0814929,0.000012519788,0.000090338945,0.000054718512,0.00015665198,0.9146847,0.0006525056,0.0024305484,0.00017431915,0.000034655517],"about_ca_topic_score_codex":0.01548116,"about_ca_topic_score_gemma":0.008821555,"teacher_disagreement_score":0.99122244,"about_ca_system_score_codex":0.0024959238,"about_ca_system_score_gemma":0.0007084646,"threshold_uncertainty_score":0.046420693},"labels":[],"label_agreement":null},{"id":"W4252966483","doi":"10.1007/s10664-021-10000-w","title":"FACER: An API usage-based code-example recommender for opportunistic reuse","year":2021,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Snippet; Java; Application programming interface; Code reuse; Android (operating system); Source code; Cluster analysis; Reuse; Information retrieval; Code (set theory); Software; World Wide Web; Data mining; Programming language; Operating system; Artificial intelligence; Set (abstract data type)","score_opus":0.11146040033542123,"score_gpt":0.3362907313551917,"score_spread":0.22483033101977049,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252966483","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32363805,0.0031377904,0.5106342,0.0016845439,0.00041722297,0.00087369734,0.012865813,0.11359651,0.033152174],"genre_scores_gemma":[0.5310368,0.00066317356,0.41542402,0.00068959605,0.00010547187,0.0003611962,0.01875448,0.001911475,0.031053754],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985934,0.0003388387,0.00007796582,0.0002626553,0.0006376084,0.00008959308],"domain_scores_gemma":[0.99485236,0.0020717331,0.00024818906,0.0015285783,0.0010155382,0.00028351892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012389029,0.0008834961,0.00094141864,0.0023194058,0.0006383525,0.0009772158,0.0018480283,0.0015010866,0.0062334975],"category_scores_gemma":[0.0104145985,0.0004253397,0.00076001993,0.0014398305,0.00020956516,0.0027800235,0.0014933329,0.0011132848,0.004971926],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009704362,0.001717135,0.04382422,0.0007863286,0.00040104034,0.00060865586,0.00062955497,0.011081082,0.015197176,0.0041123354,0.13221872,0.7884534],"study_design_scores_gemma":[0.00030861018,0.0007979674,0.023961853,0.00016096818,0.0003165022,0.0013117002,0.000539704,0.84688896,0.0185285,0.00879581,0.09814468,0.00024478493],"about_ca_topic_score_codex":0.012076695,"about_ca_topic_score_gemma":0.051054098,"teacher_disagreement_score":0.012076695,"about_ca_system_score_codex":0.00047992286,"about_ca_system_score_gemma":0.0009638646,"threshold_uncertainty_score":0.024012804},"labels":[],"label_agreement":null},{"id":"W4253058763","doi":"10.1109/mtd.2013.6608678","title":"Generating precise dependencies for large software","year":2013,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada); University of Waterloo","funders":"","keywords":"Code refactoring; Computer science; Scalability; Software engineering; Software development; Software; Source lines of code; Software system; Code (set theory); Source code; Software construction; Software evolution; Programming language; Database; Set (abstract data type)","score_opus":0.020952338323343568,"score_gpt":0.26432046299521395,"score_spread":0.24336812467187038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253058763","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07613931,0.00041642628,0.90273577,0.00025519897,0.000066173285,0.00012952207,0.0032812788,0.0146559905,0.002320391],"genre_scores_gemma":[0.28239298,0.00040183697,0.70157593,0.00011677701,0.000049996495,0.00022791096,0.009855769,0.003502559,0.0018762834],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99700016,0.00050737016,0.0002241039,0.00062608876,0.0014994723,0.00014275318],"domain_scores_gemma":[0.983008,0.008674542,0.0018526481,0.0038019817,0.0024788892,0.00018385051],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018207107,0.0010608318,0.0006759321,0.002756901,0.000820451,0.0009133901,0.0011005523,0.0008408024,0.0023107403],"category_scores_gemma":[0.020446472,0.0010113246,0.00078371953,0.0022205801,0.00062726217,0.0027325205,0.0021595675,0.0016515905,0.0009615889],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029680328,0.00017894895,0.033940148,0.001148214,0.0001539249,0.0018636036,0.0015444859,0.16299097,0.06335592,0.033721536,0.028529132,0.6722764],"study_design_scores_gemma":[0.00007717816,0.0001741813,0.019682905,0.00020492708,0.00017030726,0.0011155393,0.0003557333,0.764138,0.07494424,0.08993307,0.049053427,0.00015054496],"about_ca_topic_score_codex":0.0026371642,"about_ca_topic_score_gemma":0.0054416596,"teacher_disagreement_score":0.002756901,"about_ca_system_score_codex":0.0006244917,"about_ca_system_score_gemma":0.0018303493,"threshold_uncertainty_score":0.009628952},"labels":[],"label_agreement":null},{"id":"W4253079602","doi":"10.1002/smr.372","title":"A survey and evaluation of tool features for understanding reverse‐engineered sequence diagrams","year":2008,"lang":"en","type":"article","venue":"Journal of Software Maintenance and Evolution Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Defence Research and Development Canada; University of Victoria","funders":"","keywords":"Sequence (biology); Sequence diagram; Computer science; Set (abstract data type); Diagram; Data science; Software; Data mining; Reverse engineering; Software engineering; Unified Modeling Language; Programming language; Database","score_opus":0.23586060321562363,"score_gpt":0.39914541570243034,"score_spread":0.1632848124868067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253079602","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9404175,0.008006647,0.043155216,0.00051638874,0.000053991676,0.0006440987,0.00060386123,0.0017371171,0.004865068],"genre_scores_gemma":[0.9160282,0.0063437372,0.07371954,0.00015055781,0.000036371875,0.00049560826,0.0018674528,0.0004230951,0.0009355308],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9679693,0.01804336,0.005016636,0.0017531645,0.006408994,0.0008085236],"domain_scores_gemma":[0.6112216,0.33653125,0.013325539,0.008128829,0.02939047,0.0014023158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030559495,0.0008794797,0.00094920764,0.008995421,0.0005102419,0.0017637444,0.0011991394,0.0009695067,0.0009894035],"category_scores_gemma":[0.154243,0.00050349586,0.000661804,0.005822468,0.000548928,0.003388416,0.0012230491,0.0005326216,0.00041332064],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005893803,0.00086412084,0.10010597,0.004869003,0.00015752604,0.00039311757,0.019715965,0.0014047027,0.011850501,0.0008549336,0.0035323217,0.8556625],"study_design_scores_gemma":[0.00049428083,0.0111714015,0.64962703,0.014208906,0.0012483225,0.007860382,0.05039902,0.03152238,0.06864905,0.0031592727,0.16113718,0.0005226948],"about_ca_topic_score_codex":0.001086262,"about_ca_topic_score_gemma":0.0025480373,"teacher_disagreement_score":0.030559495,"about_ca_system_score_codex":0.00085497793,"about_ca_system_score_gemma":0.0013784919,"threshold_uncertainty_score":0.16161597},"labels":[],"label_agreement":null},{"id":"W4253188200","doi":"10.31224/osf.io/wxv9e","title":"Predicting Verification Methods from Natural Language Requirements","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Similarity (geometry); Data mining; Inference; Set (abstract data type); Natural language; Consistency (knowledge bases); Certification; Artificial intelligence; Programming language","score_opus":0.06407901670036724,"score_gpt":0.390507515203784,"score_spread":0.32642849850341676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253188200","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2969539,0.00045425698,0.69370806,0.0007818555,0.000058770165,0.0004647467,0.0023550782,0.0018658498,0.0033575178],"genre_scores_gemma":[0.70287275,0.00013073185,0.28900948,0.00013591045,0.000034690544,0.00031282427,0.0065015783,0.00015418851,0.000847764],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9942731,0.0028527149,0.00041694354,0.00085670664,0.0014251405,0.00017526334],"domain_scores_gemma":[0.9454587,0.04310294,0.0031537549,0.0017112541,0.0061882287,0.00038509813],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043859105,0.0006271214,0.0003464836,0.0034540833,0.00040342368,0.0015192536,0.0008077721,0.00094218005,0.0023177776],"category_scores_gemma":[0.038521636,0.00027806277,0.0009388068,0.0014872461,0.00062575896,0.0021842895,0.0008641314,0.001151302,0.0008609647],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005569067,0.00078004954,0.075916044,0.0010043737,0.0002067951,0.0006062589,0.0008186068,0.36528563,0.012831766,0.02310121,0.01207368,0.5068187],"study_design_scores_gemma":[0.00002045177,0.000056999703,0.0058241584,0.000037433827,0.000012728437,0.000086838525,0.00015533599,0.9793128,0.0036189263,0.009156787,0.0017008005,0.000016845031],"about_ca_topic_score_codex":0.0067594605,"about_ca_topic_score_gemma":0.008024347,"teacher_disagreement_score":0.0067594605,"about_ca_system_score_codex":0.0017329851,"about_ca_system_score_gemma":0.001685428,"threshold_uncertainty_score":0.023195148},"labels":[],"label_agreement":null},{"id":"W4253617473","doi":"10.1145/1083142.1083161","title":"SCQL","year":2005,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science","score_opus":0.015189933772932143,"score_gpt":0.2652056014566818,"score_spread":0.25001566768374966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253617473","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009090198,0.0007558247,0.61740446,0.0031315545,0.0004592097,0.0014268365,0.07924288,0.23440488,0.054084104],"genre_scores_gemma":[0.20691776,0.0019466461,0.37229824,0.0059119114,0.00060356496,0.002975972,0.3123564,0.042248223,0.054741282],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9937023,0.0011423428,0.0009822437,0.0010514265,0.0025883312,0.00053338206],"domain_scores_gemma":[0.987926,0.0042607426,0.0006089047,0.0035137103,0.003294853,0.00039585377],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065020667,0.0013761033,0.0013568708,0.0028639566,0.0016771915,0.0065917326,0.0042607365,0.002134304,0.038013298],"category_scores_gemma":[0.019603772,0.0010849377,0.0021585762,0.0037002603,0.0014200738,0.009644809,0.0055459184,0.0026433293,0.017166577],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013426492,0.000275934,0.006267356,0.0020717683,0.0002782299,0.0006957609,0.00154608,0.012115507,0.007681683,0.25069332,0.53440964,0.1826221],"study_design_scores_gemma":[0.00031802963,0.00012692023,0.0013160482,0.00028653487,0.00009361171,0.00057323946,0.00081183104,0.089874335,0.0109933475,0.17210306,0.72332984,0.00017324479],"about_ca_topic_score_codex":0.015608845,"about_ca_topic_score_gemma":0.010981167,"teacher_disagreement_score":0.038013298,"about_ca_system_score_codex":0.0028936863,"about_ca_system_score_gemma":0.0037990026,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4253675410","doi":"10.1109/iwpse.2004.1334769","title":"Evolution spectrographs: visualizing punctuated change in software evolution","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Software evolution; Punctuated equilibrium; Computer science; Punctuation; Software; Visualization; Software visualization; Spectrograph; Process (computing); Software development; Paleontology; Data mining; Artificial intelligence; Geology; Software construction; Programming language; Astronomy; Physics","score_opus":0.025745816601175902,"score_gpt":0.28801124415715257,"score_spread":0.26226542755597665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253675410","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5408054,0.0011260976,0.41448176,0.0011224201,0.00015568797,0.000283154,0.0061344653,0.02547683,0.010414266],"genre_scores_gemma":[0.77718955,0.0006251465,0.21638,0.00014938346,0.00006445616,0.00018257536,0.0021108366,0.0009264526,0.0023716257],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99976414,0.00006807974,0.00001576606,0.000040154675,0.00007694613,0.00003492297],"domain_scores_gemma":[0.9979342,0.0010235102,0.00036037352,0.00014784996,0.0003622672,0.00017179249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00064192095,0.00047953028,0.00020677173,0.0051244535,0.00039027614,0.0007811193,0.00039887073,0.000730463,0.0028489125],"category_scores_gemma":[0.003303438,0.00022464564,0.00028225573,0.002536034,0.00030173475,0.0012914385,0.0010086977,0.0007687397,0.00030134176],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015277786,0.0003658001,0.057921004,0.0010721423,0.00027830582,0.0017479254,0.01646304,0.026649062,0.2060038,0.019985465,0.032957856,0.6350277],"study_design_scores_gemma":[0.00022341844,0.000579337,0.2892632,0.00045473283,0.00024700584,0.0031199001,0.005064924,0.45992497,0.119408675,0.028356615,0.09291438,0.00044293588],"about_ca_topic_score_codex":0.0044244463,"about_ca_topic_score_gemma":0.00491584,"teacher_disagreement_score":0.0051244535,"about_ca_system_score_codex":0.0003864506,"about_ca_system_score_gemma":0.00041612066,"threshold_uncertainty_score":0.009530604},"labels":[],"label_agreement":null},{"id":"W4253845945","doi":"10.22215/etd/2021-14475","title":"Understanding and Predicting Software Developer Expertise in Stack Overflow and GitHub","year":2021,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Software; Data science; Coding (social sciences); Software engineering; Knowledge management; Operating system","score_opus":0.05876980672213129,"score_gpt":0.28305342837367004,"score_spread":0.22428362165153876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253845945","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99755114,0.00013309148,0.0006199723,0.00012639156,0.0000030357085,0.000011234345,0.00006241395,0.000017603219,0.0014750962],"genre_scores_gemma":[0.99764293,0.00014928864,0.0012734929,0.00002211678,0.000005792983,0.00001258269,0.0001832583,0.0000139921785,0.0006964964],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99802655,0.0005916261,0.00015457292,0.00038040205,0.00055976596,0.00028708874],"domain_scores_gemma":[0.9612667,0.024828171,0.007207949,0.00082845456,0.0031207637,0.0027479907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031413722,0.00039433816,0.00030305804,0.0032119085,0.0006764657,0.0017446406,0.00043802388,0.0008054613,0.0015825875],"category_scores_gemma":[0.038166914,0.00027522742,0.0002496811,0.0015991583,0.0005256144,0.0034633726,0.0017353294,0.0007762652,0.0005742512],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013860881,0.00023569474,0.9301347,0.000110386914,0.00003836095,0.00048114476,0.01609961,0.0018661596,0.00095177535,0.00044349316,0.0016359054,0.04786398],"study_design_scores_gemma":[0.000009446135,0.00022291532,0.95124424,0.00015238936,0.000037350714,0.0005508792,0.020824712,0.021160368,0.0012241545,0.0012500593,0.0032763444,0.000047098347],"about_ca_topic_score_codex":0.0106069865,"about_ca_topic_score_gemma":0.018882763,"teacher_disagreement_score":0.0106069865,"about_ca_system_score_codex":0.0008378776,"about_ca_system_score_gemma":0.0007347899,"threshold_uncertainty_score":0.021090448},"labels":[],"label_agreement":null},{"id":"W4253901503","doi":"10.32920/ryerson.14648622.v1","title":"On Empirically Examining The Effectiveness Of Deep Learning-Based Bug Localization Models","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software bug; Computer science; Java; Set (abstract data type); Artificial intelligence; Convolutional neural network; Software; Baseline (sea); Machine learning; Deep learning; Convolution (computer science); Software engineering; State (computer science); Artificial neural network; Programming language","score_opus":0.033230855442043704,"score_gpt":0.2871205612570827,"score_spread":0.253889705815039,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253901503","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9066384,0.008643059,0.0666111,0.003691472,0.0003315026,0.00025496355,0.0021807505,0.0021593135,0.009489511],"genre_scores_gemma":[0.97249776,0.001058041,0.02262434,0.00037232458,0.00006445728,0.00009948421,0.0022163386,0.00009136518,0.00097589183],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99542266,0.0021925499,0.0004475129,0.0008661835,0.0007376974,0.00033332955],"domain_scores_gemma":[0.9391235,0.04854934,0.0029101355,0.003895964,0.004634511,0.00088658824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012441448,0.0031701184,0.0013602679,0.0020496568,0.00084451283,0.0021195032,0.0022226898,0.003560991,0.0018640254],"category_scores_gemma":[0.060674608,0.00088976504,0.0009982137,0.0018733937,0.0015750283,0.006096322,0.0017745451,0.0036720482,0.0007127594],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014326922,0.0012363022,0.04046419,0.0008767899,0.0006226913,0.00014404485,0.00018573357,0.84326476,0.0022521361,0.0031286576,0.0070707765,0.09932121],"study_design_scores_gemma":[0.00008840155,0.00067299936,0.0035010257,0.000146479,0.00014844288,0.00006563409,0.000096204734,0.9891777,0.0024119809,0.002928714,0.00073075824,0.000031628108],"about_ca_topic_score_codex":0.021703608,"about_ca_topic_score_gemma":0.021387901,"teacher_disagreement_score":0.021703608,"about_ca_system_score_codex":0.0031580415,"about_ca_system_score_gemma":0.0018400372,"threshold_uncertainty_score":0.06579739},"labels":[],"label_agreement":null},{"id":"W4254317839","doi":"10.1145/585064.585065","title":"The relevance of software documentation, tools and technologies","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Documentation; Software documentation; Software engineering; Relevance (law); Computer science; Internal documentation; Technical documentation; Software; Software maintenance; Software development; Knowledge management; World Wide Web; Data science; Software construction; Programming language","score_opus":0.021083646717175407,"score_gpt":0.25959428714179356,"score_spread":0.23851064042461814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4254317839","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9808456,0.0023369663,0.0025749004,0.0022959404,0.00003581867,0.000055720095,0.00005506768,0.000031430827,0.01176854],"genre_scores_gemma":[0.99634343,0.00076829595,0.001568668,0.00039320736,0.000037671718,0.00002239788,0.000059621372,0.000015282827,0.0007915598],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.96859276,0.014967444,0.0029736285,0.0009769378,0.011331924,0.0011574501],"domain_scores_gemma":[0.83256876,0.11908057,0.021545807,0.0038211525,0.018003095,0.0049805893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017651003,0.0002166308,0.0003890692,0.0034065358,0.0017528214,0.0038649696,0.0004020536,0.0012300957,0.0018899019],"category_scores_gemma":[0.14672597,0.00031975732,0.0003298576,0.0025003501,0.0015602123,0.004125872,0.002356615,0.0013119702,0.00035558402],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004602269,0.00029198616,0.61217767,0.0017467596,0.00015294642,0.001201517,0.11990836,0.00058910315,0.0068821986,0.00412616,0.0031712533,0.24929191],"study_design_scores_gemma":[0.000051479186,0.0012732719,0.7485338,0.0016245833,0.00018493852,0.004444783,0.18630567,0.0017502623,0.002312018,0.006758057,0.0466049,0.00015619896],"about_ca_topic_score_codex":0.0014977582,"about_ca_topic_score_gemma":0.00238422,"teacher_disagreement_score":0.017651003,"about_ca_system_score_codex":0.0013027982,"about_ca_system_score_gemma":0.0021315739,"threshold_uncertainty_score":0.09334856},"labels":[],"label_agreement":null},{"id":"W4254783527","doi":"10.1109/vissoft.2016.18","title":"Merge-Tree: Visualizing the Integration of Commits into Linux","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Commit; Computer science; Linux kernel; Merge (version control); Operating system; Programming language; Database; Parallel computing","score_opus":0.03056849072648814,"score_gpt":0.3205195445459963,"score_spread":0.28995105381950814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4254783527","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12813632,0.0026250773,0.5820397,0.0023262252,0.00084450806,0.0005802254,0.04171874,0.21927346,0.022455715],"genre_scores_gemma":[0.41806954,0.0018643183,0.5138144,0.00049093604,0.00015315104,0.0004857684,0.03820473,0.018280637,0.008636479],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994974,0.00009324765,0.000050627517,0.00010045734,0.0001765531,0.00008170287],"domain_scores_gemma":[0.99736196,0.0011014234,0.00031662526,0.00036024777,0.00050879735,0.00035102412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010885508,0.0011643919,0.0004564771,0.004357948,0.0008363389,0.001972027,0.00140891,0.00088137225,0.008543571],"category_scores_gemma":[0.005606258,0.0005009369,0.00087082945,0.0033767056,0.00050585176,0.0029885648,0.002535536,0.0017556018,0.0019980273],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002394226,0.00060885097,0.062115133,0.0027506168,0.00035311325,0.0025759763,0.021376157,0.048080903,0.028410228,0.04924447,0.29328427,0.48880607],"study_design_scores_gemma":[0.00036554373,0.00038433584,0.047263887,0.00081694266,0.00024395602,0.0019675358,0.0060363836,0.36756802,0.034014117,0.06746739,0.473425,0.00044693187],"about_ca_topic_score_codex":0.01349789,"about_ca_topic_score_gemma":0.017126206,"teacher_disagreement_score":0.01349789,"about_ca_system_score_codex":0.00068067596,"about_ca_system_score_gemma":0.0016431748,"threshold_uncertainty_score":0.028581142},"labels":[],"label_agreement":null},{"id":"W4254933085","doi":"10.4018/978-1-4666-2470-2.ch002","title":"FTT","year":2012,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Ajax; Code refactoring; Computer science; Web application; World Wide Web; Web modeling; Web page; Software engineering; Programming language; Software","score_opus":0.023720967664058887,"score_gpt":0.25850403363190627,"score_spread":0.23478306596784737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4254933085","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017860517,0.005024092,0.052072357,0.0017608248,0.002162411,0.00015329762,0.002375808,0.0066069737,0.92805827],"genre_scores_gemma":[0.008907784,0.003298877,0.01919417,0.00056588714,0.00037119925,0.00009292884,0.0024468657,0.0016803988,0.9634419],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99947745,0.00005185096,0.000031048643,0.00010243683,0.00029414677,0.000043056647],"domain_scores_gemma":[0.99928313,0.00020878908,0.00003287865,0.00017591531,0.00021547674,0.00008381976],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00043086195,0.0007952046,0.00041199906,0.0026159587,0.0010170024,0.0028953694,0.0015277403,0.0017871046,0.30889323],"category_scores_gemma":[0.0018607937,0.00032734455,0.0007087759,0.0024469423,0.000591121,0.0036423753,0.0016467543,0.0015149406,0.15024878],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031427207,0.00004691092,0.00012750957,0.00036449157,0.000004908902,0.0003167711,0.00047007995,0.00035274276,0.0021964798,0.07285907,0.3596827,0.5635468],"study_design_scores_gemma":[0.0000029667392,0.000008462907,0.00008052073,0.00005689478,0.0000016092898,0.00037638537,0.000025254056,0.0001542423,0.00040936552,0.0050503807,0.9938292,0.0000047631693],"about_ca_topic_score_codex":0.0018608043,"about_ca_topic_score_gemma":0.0026959714,"teacher_disagreement_score":0.6911068,"about_ca_system_score_codex":0.0013587491,"about_ca_system_score_gemma":0.0010050294,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4255017042","doi":"10.22215/etd/2009-09017","title":"The relationship between code based object oriented system design metrics and product quality","year":2009,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Library and Archives Canada","funders":"","keywords":"Product (mathematics); Computer science; Quality (philosophy); Code (set theory); Object (grammar); Programming language; Mathematics; Artificial intelligence; Philosophy; Epistemology","score_opus":0.09418904932227146,"score_gpt":0.3514827006747485,"score_spread":0.25729365135247706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255017042","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8692712,0.002551461,0.11132518,0.0016928163,0.00009622632,0.00008957039,0.00050541945,0.00044668323,0.014021445],"genre_scores_gemma":[0.9851746,0.00028108293,0.013118552,0.000051614254,0.000026808575,0.000030432493,0.00027281034,0.0000764052,0.00096771476],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9893952,0.003955676,0.0008450651,0.0005807887,0.004971066,0.0002521788],"domain_scores_gemma":[0.71326977,0.20540734,0.03977357,0.006838921,0.03254943,0.0021609892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007024287,0.00036784826,0.00023158791,0.004129981,0.00023110771,0.0026559087,0.00039288172,0.0007974367,0.0014456882],"category_scores_gemma":[0.14764337,0.00032013343,0.00022128278,0.003724845,0.00066907145,0.003005537,0.00086486497,0.0009407879,0.00030027778],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048346265,0.00026070236,0.6780914,0.00038449766,0.00044450702,0.00016650536,0.0017317897,0.025669085,0.008572187,0.026245283,0.0030671686,0.25488335],"study_design_scores_gemma":[0.00004909836,0.001003194,0.8088926,0.00019850145,0.00025323723,0.000567367,0.0008659643,0.13845886,0.009571549,0.033105608,0.0069198236,0.00011415733],"about_ca_topic_score_codex":0.0016171109,"about_ca_topic_score_gemma":0.0019888973,"teacher_disagreement_score":0.007024287,"about_ca_system_score_codex":0.00097654527,"about_ca_system_score_gemma":0.00073989085,"threshold_uncertainty_score":0.037148416},"labels":[],"label_agreement":null},{"id":"W4255080050","doi":"10.4018/9781591408963.ch005","title":"Design Patterns as Laws of Quality","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Quality (philosophy); Computer science; Software quality; Software; Object (grammar); Measure (data warehouse); Software quality control; Data mining; Software engineering; Artificial intelligence; Software development; Programming language","score_opus":0.06916189357925827,"score_gpt":0.302254467090375,"score_spread":0.23309257351111673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255080050","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026416695,0.0088724,0.62755704,0.017535264,0.00045498542,0.0003118784,0.00024440704,0.0008502651,0.3177571],"genre_scores_gemma":[0.48169065,0.009803868,0.4322953,0.0029709013,0.00044650558,0.0015722363,0.00042907152,0.0007370343,0.07005439],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9948167,0.0021117243,0.00044481124,0.00074819836,0.0015962028,0.00028232],"domain_scores_gemma":[0.9901862,0.005371299,0.0007864818,0.002411298,0.0009873031,0.00025739134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050329175,0.00080157927,0.000584037,0.0021338144,0.0017357557,0.0071519525,0.0015762526,0.0028576574,0.006220036],"category_scores_gemma":[0.014373402,0.0011461049,0.0009768109,0.0022410066,0.014931746,0.011905684,0.0029084506,0.003733795,0.0014356086],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000025595446,0.0000067746487,0.00014425971,0.000052205673,0.000004797668,0.000025471078,0.00058891374,0.0011758415,0.00011007604,0.98706275,0.0012978109,0.00952865],"study_design_scores_gemma":[0.0000078847315,0.000006574063,0.00008440358,0.00006122633,0.0000047939157,0.000052255407,0.00009292393,0.0021098733,0.00014907135,0.9692021,0.028224094,0.0000048661686],"about_ca_topic_score_codex":0.0020812815,"about_ca_topic_score_gemma":0.0016576169,"teacher_disagreement_score":0.0071519525,"about_ca_system_score_codex":0.0040639434,"about_ca_system_score_gemma":0.00207656,"threshold_uncertainty_score":0.02948618},"labels":[],"label_agreement":null},{"id":"W4255420728","doi":"10.1109/icse.2013.6606693","title":"Using mutation analysis for a model-clone detector comparison framework","year":2013,"lang":"en","type":"article","venue":"2013 35th International Conference on Software Engineering (ICSE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"clone (Java method); Computer science; Mutation; Precision and recall; Data mining; Artificial intelligence; Genetics; Biology; Gene","score_opus":0.09241037752275671,"score_gpt":0.34181581097862745,"score_spread":0.24940543345587074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255420728","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046368022,0.000045345787,0.9929751,0.00009584637,0.000015471844,0.00010008231,0.000028725488,0.0016543767,0.00044827137],"genre_scores_gemma":[0.09536165,0.00005559321,0.90338993,0.00010075965,0.000013672185,0.00016637666,0.0001256349,0.00037490093,0.00041145156],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9812251,0.005718878,0.0017621829,0.0024735779,0.007934049,0.0008862485],"domain_scores_gemma":[0.9723022,0.013507872,0.00345535,0.0041805627,0.0061217146,0.00043235163],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019643439,0.0017196726,0.0018598514,0.006426702,0.0014048986,0.005346346,0.00435879,0.0024244054,0.0018743585],"category_scores_gemma":[0.053525906,0.00092007115,0.0030908275,0.0025191528,0.0028392298,0.0061835796,0.0038939924,0.0030082276,0.00046669744],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004964562,0.0005983678,0.014191326,0.0006457703,0.00053169864,0.0009970804,0.0013490588,0.27074608,0.042834777,0.3018057,0.003115517,0.36268818],"study_design_scores_gemma":[0.000064995256,0.00035716486,0.0014420887,0.00015999426,0.00020969183,0.0005491891,0.00022972102,0.86577713,0.03859552,0.08236009,0.010108503,0.00014606149],"about_ca_topic_score_codex":0.0058159325,"about_ca_topic_score_gemma":0.004631905,"teacher_disagreement_score":0.019643439,"about_ca_system_score_codex":0.0038333372,"about_ca_system_score_gemma":0.0059687058,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4255597738","doi":"10.22215/etd/2007-08350","title":"Using machine learning to support software debugging","year":2007,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Canadian Heritage","funders":"","keywords":"Debugging; Computer science; Software; Software engineering; Operating system; Programming language","score_opus":0.03582681888709189,"score_gpt":0.3505824444857276,"score_spread":0.31475562559863574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255597738","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06965471,0.00068593846,0.90345037,0.0011942068,0.00018587054,0.00012433072,0.00019818298,0.017523104,0.0069832797],"genre_scores_gemma":[0.44021723,0.00045126418,0.55315924,0.00018616114,0.000086095955,0.00010936007,0.00075597654,0.00051345007,0.0045212684],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985758,0.0005740531,0.00010557687,0.00027493256,0.00038371418,0.00008598348],"domain_scores_gemma":[0.98878634,0.006958244,0.00089390064,0.0013513449,0.001733198,0.00027696835],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020864338,0.00065238675,0.00046039847,0.0019005501,0.000494303,0.0014271628,0.001089539,0.00084020133,0.0022649576],"category_scores_gemma":[0.01667101,0.000453453,0.0004328524,0.0009115579,0.00039447314,0.0029056391,0.0007768869,0.0013412156,0.0013227989],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023783451,0.0005134035,0.007215508,0.00015237415,0.00008306828,0.00013350778,0.00021754956,0.08871355,0.00835943,0.005644526,0.0074471086,0.8812822],"study_design_scores_gemma":[0.000028176402,0.000056954766,0.00077834836,0.000030431776,0.000021873955,0.000047964793,0.000023767969,0.9728127,0.010000694,0.012139382,0.004044057,0.000015623908],"about_ca_topic_score_codex":0.0024747788,"about_ca_topic_score_gemma":0.0047517056,"teacher_disagreement_score":0.0024747788,"about_ca_system_score_codex":0.00060362276,"about_ca_system_score_gemma":0.0010525946,"threshold_uncertainty_score":0.01103425},"labels":[],"label_agreement":null},{"id":"W4255603372","doi":"10.1109/tools.1997.654721","title":"Mapping the OO-Jacobson approach into function point analysis","year":2002,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Function point; Computer science; Function (biology); Point (geometry); Software; Object (grammar); Set (abstract data type); Algorithm; Software development; Theoretical computer science; Data mining; Programming language; Mathematics; Artificial intelligence","score_opus":0.038177343865060734,"score_gpt":0.22923606418666426,"score_spread":0.19105872032160354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255603372","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005288509,0.0003895573,0.97773874,0.0004924179,0.000048570433,0.00006244721,0.00002386364,0.00025828907,0.015697582],"genre_scores_gemma":[0.1895762,0.0013777155,0.79767674,0.00051177386,0.00015704913,0.00030544432,0.00012432948,0.00040563542,0.009865245],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993149,0.0019414049,0.00043066134,0.000889904,0.003352098,0.00023694363],"domain_scores_gemma":[0.9917508,0.0036614584,0.00062261947,0.0017163439,0.0020657019,0.0001832247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039948067,0.0009962446,0.0006996028,0.0047750976,0.00069106027,0.003746319,0.0015756111,0.0010325345,0.0023418453],"category_scores_gemma":[0.011273629,0.00071460975,0.0011489082,0.0024795237,0.004749246,0.0047265743,0.0024570755,0.0021703513,0.00080573803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019728677,0.000051218387,0.0018820858,0.00019664054,0.0000313219,0.00010054399,0.0013876479,0.00834248,0.0025318735,0.81992865,0.0011868768,0.1643409],"study_design_scores_gemma":[0.000017502718,0.000083042905,0.0026793273,0.00013491257,0.000034022087,0.0003028798,0.0005515864,0.08885695,0.0042036204,0.84585714,0.057226155,0.000052738436],"about_ca_topic_score_codex":0.0045830174,"about_ca_topic_score_gemma":0.0021209538,"teacher_disagreement_score":0.0047750976,"about_ca_system_score_codex":0.0034660117,"about_ca_system_score_gemma":0.0026467573,"threshold_uncertainty_score":0.025147796},"labels":[],"label_agreement":null},{"id":"W4255649397","doi":"10.32920/ryerson.14638470","title":"Partially observable Markov Decision Process to prioritize software defects","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Dependency (UML); Partially observable Markov decision process; Dependency graph; Exploit; Software bug; Prioritization; Software; Graph; Process (computing); Software quality; Software regression; Data mining; Markov chain; Markov model; Software development; Machine learning; Artificial intelligence; Computer security; Engineering; Theoretical computer science","score_opus":0.025687265362385556,"score_gpt":0.2996946730810493,"score_spread":0.2740074077186637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255649397","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0643449,0.00037279818,0.9299793,0.0007130605,0.0000981587,0.000280298,0.00047093353,0.00054997805,0.0031904853],"genre_scores_gemma":[0.90397125,0.0003040143,0.09211806,0.00021079704,0.00005233215,0.00047528293,0.0004958757,0.000045795387,0.0023266047],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99723166,0.0010189217,0.00014005158,0.0006196949,0.00056309334,0.00042651143],"domain_scores_gemma":[0.98432195,0.012832988,0.0011563625,0.00022832277,0.0010211553,0.00043917203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032714957,0.0016301847,0.001869381,0.0012275951,0.0006799754,0.0014091147,0.0014940625,0.0014450949,0.0037938212],"category_scores_gemma":[0.010737421,0.0008337819,0.0015206502,0.0009363327,0.0014096837,0.001389338,0.0012169895,0.0024517856,0.00028911428],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016361738,0.00006622217,0.002365075,0.00013392106,0.00006179725,0.00017434313,0.000087197055,0.97409874,0.0005330836,0.01298065,0.0004959079,0.008839363],"study_design_scores_gemma":[0.00002480386,0.000036371486,0.00024555004,0.0000096395615,0.000019954596,0.000011596611,0.000011662201,0.9936441,0.00015738598,0.005691555,0.00013954587,0.000007800476],"about_ca_topic_score_codex":0.02224345,"about_ca_topic_score_gemma":0.01455428,"teacher_disagreement_score":0.02224345,"about_ca_system_score_codex":0.0022320237,"about_ca_system_score_gemma":0.0035682977,"threshold_uncertainty_score":0.044227958},"labels":[],"label_agreement":null},{"id":"W4255771497","doi":"10.1109/icse.2015.97","title":"Discovering Information Explaining API Types Using Text Classification","year":2015,"lang":"en","type":"article","venue":"2015 IEEE/ACM 37th IEEE International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Application programming interface; Interface (matter); Information retrieval; Precision and recall; Software; Recall; Artificial intelligence; Natural language processing; Programming language","score_opus":0.12655072206829188,"score_gpt":0.33994261545711246,"score_spread":0.21339189338882059,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255771497","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43151635,0.005014271,0.49135086,0.001407573,0.00052905397,0.001405846,0.027976047,0.027299205,0.0135008115],"genre_scores_gemma":[0.5619267,0.0015290803,0.37289453,0.00034372136,0.00054484524,0.0008817978,0.053478822,0.0008126842,0.007587837],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99868685,0.00023150288,0.00017171624,0.00043605905,0.00037870405,0.00009520542],"domain_scores_gemma":[0.9866361,0.007800018,0.0014907884,0.0007971105,0.0029720112,0.00030392088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014922261,0.0019103804,0.00076107733,0.012256079,0.0007178009,0.0015922633,0.0011218047,0.0013549344,0.003356591],"category_scores_gemma":[0.01180577,0.00033047836,0.001281734,0.004551032,0.00033254657,0.0033725982,0.00087109505,0.001312829,0.003370908],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005324176,0.0007263599,0.034398373,0.0010358251,0.00014302492,0.000826984,0.0009163931,0.0064793304,0.037200823,0.0016833605,0.034824476,0.8812327],"study_design_scores_gemma":[0.00020884973,0.00097348366,0.10218293,0.0008235319,0.0007487753,0.002192398,0.002456477,0.68022764,0.097528994,0.021670194,0.09071943,0.00026729793],"about_ca_topic_score_codex":0.003435685,"about_ca_topic_score_gemma":0.004005061,"teacher_disagreement_score":0.012256079,"about_ca_system_score_codex":0.0007423915,"about_ca_system_score_gemma":0.0010594684,"threshold_uncertainty_score":0.011228979},"labels":[],"label_agreement":null},{"id":"W4255823579","doi":"10.4018/978-1-60566-060-8.ch055","title":"Constructivist Learning During Software Development","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Laurentian University","funders":"","keywords":"Constructivist teaching methods; Programmer; Denial; Computer science; Process (computing); Documentation; Knowledge management; Cognitive science; Mathematics education; Psychology; Programming language; Teaching method","score_opus":0.013062628960595535,"score_gpt":0.23222253365044973,"score_spread":0.2191599046898542,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255823579","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031857014,0.03803124,0.37540445,0.016229028,0.00046963798,0.00012185066,0.00006062183,0.00033028098,0.53749585],"genre_scores_gemma":[0.64652175,0.03330322,0.16364855,0.0020551481,0.00044227598,0.00036130648,0.00015120674,0.00027696227,0.15323949],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9986595,0.0007564313,0.00003538038,0.00015139609,0.00033038252,0.00006697498],"domain_scores_gemma":[0.9976145,0.0020056278,0.00006932021,0.00017138616,0.000091075475,0.00004805693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022824262,0.0005695647,0.00025762813,0.00076862023,0.0009133579,0.0029816818,0.0012984872,0.001476052,0.0029632684],"category_scores_gemma":[0.003093482,0.00035320365,0.00028999438,0.000941029,0.012228078,0.005033973,0.0021404375,0.0038326492,0.00072658894],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012884127,0.000028326092,0.00023704713,0.00023892132,0.000006916762,0.000112474765,0.012659586,0.0025746196,0.0003967438,0.911256,0.0038734612,0.0686031],"study_design_scores_gemma":[0.00001822508,0.000036138696,0.000323727,0.0003910847,0.000006433544,0.00020648188,0.0012605378,0.0033255864,0.00086856453,0.8037313,0.18981798,0.000013959774],"about_ca_topic_score_codex":0.0016713493,"about_ca_topic_score_gemma":0.0023817634,"teacher_disagreement_score":0.0041327965,"about_ca_system_score_codex":0.0041327965,"about_ca_system_score_gemma":0.0020350183,"threshold_uncertainty_score":0.029985666},"labels":[],"label_agreement":null},{"id":"W4256027430","doi":"10.1109/icse.2013.6606748","title":"Identifying failure inducing developer pairs within developer networks","year":2013,"lang":"en","type":"article","venue":"2013 35th International Conference on Software Engineering (ICSE)","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software engineering; Code (set theory); Software; Liberian dollar; Software bug; Software maintenance; Software system; Software development; Source code; Software construction; Programming language","score_opus":0.04027914682502887,"score_gpt":0.2656480069125179,"score_spread":0.22536886008748902,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4256027430","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8918222,0.0005888838,0.0980593,0.00043816422,0.000044704084,0.0002896873,0.0008356424,0.00020714849,0.0077143013],"genre_scores_gemma":[0.96521765,0.00030264514,0.03072435,0.000058369962,0.00002739617,0.00022082293,0.0011372223,0.000053347598,0.0022581858],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99532133,0.0016152383,0.00021912117,0.0012814697,0.0010692243,0.0004936451],"domain_scores_gemma":[0.9708811,0.016451418,0.0051337443,0.002135545,0.0037037001,0.0016943924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002985736,0.0006495812,0.0006230491,0.006164845,0.0017047176,0.0015367834,0.001069786,0.0012496088,0.0021607268],"category_scores_gemma":[0.03466648,0.0006447306,0.00042732502,0.0028155874,0.0008159536,0.0026774602,0.0027244163,0.0008385565,0.00080827996],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004312767,0.0002852482,0.8452074,0.00023947685,0.00014787563,0.0029315182,0.0069708796,0.018981837,0.008092144,0.018978998,0.0043975413,0.09333587],"study_design_scores_gemma":[0.000111945395,0.0005067811,0.39939266,0.0002603563,0.00048273252,0.007939663,0.01623591,0.44613484,0.015922,0.07569604,0.03715262,0.00016443722],"about_ca_topic_score_codex":0.0050761653,"about_ca_topic_score_gemma":0.008883421,"teacher_disagreement_score":0.006164845,"about_ca_system_score_codex":0.0012425226,"about_ca_system_score_gemma":0.00114245,"threshold_uncertainty_score":0.015790224},"labels":[],"label_agreement":null},{"id":"W4256061005","doi":"10.32920/ryerson.14665455","title":"Speeding up calibration of latent Dirichlet allocation model to improve topic analysis in software engineering","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Topic model; Exploit; Software; Simple (philosophy); Dirichlet distribution; Information retrieval; Data mining; Data science; Theoretical computer science; Mathematics; Programming language","score_opus":0.027622280103591152,"score_gpt":0.27322675240945776,"score_spread":0.2456044723058666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4256061005","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016074697,0.0004281268,0.9811916,0.0003521687,0.00007010012,0.00007861936,0.00010127336,0.0010220807,0.00068142585],"genre_scores_gemma":[0.2628559,0.0006363863,0.7319005,0.00034434898,0.00019912483,0.0006269001,0.0010826963,0.00066886825,0.0016852921],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9928925,0.004406427,0.00036641918,0.0012625335,0.00075831474,0.00031377963],"domain_scores_gemma":[0.97173667,0.0223188,0.00083304266,0.002463382,0.0023113696,0.00033677102],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011431054,0.001254036,0.0015961614,0.0031625964,0.0014487596,0.0031465702,0.0019021716,0.0021739078,0.0031682358],"category_scores_gemma":[0.06728162,0.0010946462,0.0015153886,0.0028709115,0.0012158736,0.0061109015,0.0040414217,0.0045800866,0.0023945302],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047744322,0.0003240684,0.010789215,0.0004557642,0.00025367012,0.00019449824,0.0026389915,0.41778693,0.011137599,0.058802728,0.01017349,0.4869657],"study_design_scores_gemma":[0.000036798996,0.00003084851,0.00095021335,0.000035988844,0.000022957398,0.000050861523,0.0001450951,0.9486023,0.002643851,0.044681683,0.002757736,0.000041715513],"about_ca_topic_score_codex":0.0068784696,"about_ca_topic_score_gemma":0.00745724,"teacher_disagreement_score":0.011431054,"about_ca_system_score_codex":0.0015563731,"about_ca_system_score_gemma":0.0025388482,"threshold_uncertainty_score":0.06045389},"labels":[],"label_agreement":null},{"id":"W4256378361","doi":"10.1109/csmr.1999.756690","title":"A change impact model for changeability assessment in object-oriented software systems","year":2003,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Maintainability; Software maintenance; Computer science; Software metric; Software quality; Change impact analysis; Software system; Software; Software engineering; Object-oriented programming; Software sizing; Reliability engineering; Software development; Conceptual model; Object-oriented design; Systems engineering; Software construction; Engineering; Programming language; Database","score_opus":0.06443102988361787,"score_gpt":0.3467746561491911,"score_spread":0.28234362626557324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4256378361","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04067639,0.00033904682,0.94713134,0.0005994765,0.00006619309,0.00034209742,0.00016732223,0.00059315836,0.010084981],"genre_scores_gemma":[0.82046044,0.00046585372,0.17143258,0.00027471432,0.0001074056,0.00088138,0.0003730546,0.00022516777,0.005779454],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99583286,0.0015628928,0.00022521363,0.00037630336,0.0017709197,0.00023176115],"domain_scores_gemma":[0.99024105,0.0065435357,0.0011400793,0.0007703238,0.0010693102,0.00023572818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032804548,0.0011800594,0.0006065209,0.0035565693,0.000519756,0.0017310288,0.0014610918,0.0021258106,0.0025898828],"category_scores_gemma":[0.019584958,0.00042670098,0.0011798341,0.0020457795,0.0017298649,0.0037912056,0.0011637599,0.0018049885,0.00068997504],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015556016,0.00037614236,0.00867054,0.00025666386,0.00012732002,0.00041782533,0.0006786944,0.7574909,0.007022079,0.15165529,0.0016584296,0.07149051],"study_design_scores_gemma":[0.000024831568,0.00024030644,0.0026813133,0.000030123778,0.000042693508,0.00018423669,0.00007865879,0.92838776,0.0012436351,0.06392566,0.0031285353,0.000032253836],"about_ca_topic_score_codex":0.0034392574,"about_ca_topic_score_gemma":0.0019916687,"teacher_disagreement_score":0.0035565693,"about_ca_system_score_codex":0.0018528935,"about_ca_system_score_gemma":0.00082720845,"threshold_uncertainty_score":0.017348945},"labels":[],"label_agreement":null},{"id":"W4256503404","doi":"10.32920/ryerson.14646261.v1","title":"On Predicting Rediscoveries of Software Defects","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Eclipse; Downtime; Computer science; Software bug; Forcing (mathematics); Software; Customer satisfaction; Quality (philosophy); Software quality; Software engineering; Data science; Software development; Operating system; Business; Marketing","score_opus":0.018148737736127903,"score_gpt":0.26380164501338726,"score_spread":0.24565290727725936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4256503404","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95857227,0.0026182171,0.025302552,0.0016206756,0.0000904725,0.000118086195,0.008348831,0.0013014411,0.0020273898],"genre_scores_gemma":[0.96301633,0.0008078575,0.018566966,0.00017775816,0.00009101616,0.00006304697,0.015874976,0.000048754297,0.001353292],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984861,0.00033460037,0.00013595527,0.0005418895,0.00034198485,0.00015952412],"domain_scores_gemma":[0.9804176,0.012965155,0.0024688118,0.0012927646,0.0022425002,0.0006130719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031539071,0.0011415357,0.0007467812,0.0053547486,0.0004795182,0.0012319373,0.0013241142,0.0019212398,0.00082984107],"category_scores_gemma":[0.01764614,0.0004871106,0.00089956645,0.0037812553,0.0003548192,0.001949881,0.00073405117,0.0017675346,0.0008166902],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033592858,0.00080561946,0.7796904,0.00027150844,0.00031489463,0.00040289218,0.0003394501,0.086298294,0.0012454216,0.0008752359,0.011109525,0.11831078],"study_design_scores_gemma":[0.00004260467,0.00026440967,0.18706237,0.00009020794,0.00017014684,0.00040722208,0.00036177223,0.80493426,0.0012940695,0.001813078,0.0035071443,0.000052784235],"about_ca_topic_score_codex":0.049042627,"about_ca_topic_score_gemma":0.06764831,"teacher_disagreement_score":0.049042627,"about_ca_system_score_codex":0.0008776104,"about_ca_system_score_gemma":0.0009606993,"threshold_uncertainty_score":0.09751439},"labels":[],"label_agreement":null},{"id":"W4256657178","doi":"10.1109/msr.2015.23","title":"Investigating Code Review Practices in Defective Files: An Empirical Study of the Qt System","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Code review; Computer science; Source code; Documentation; Software quality; Software bug; Code (set theory); Empirical research; Software; Software engineering; Evolvability; Software development; Programming language; Set (abstract data type)","score_opus":0.1758062651429521,"score_gpt":0.4161226211492933,"score_spread":0.24031635600634121,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4256657178","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99918026,0.0000871379,0.00032386865,0.000077710836,0.0000020875184,0.00003088778,0.000053101998,0.00001021007,0.00023459144],"genre_scores_gemma":[0.9985684,0.00011804413,0.0007134666,0.00005600135,0.000007579754,0.00004053791,0.00015593674,0.000020973808,0.00031891195],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98445916,0.0055751274,0.0014725111,0.0019372766,0.00579674,0.0007592401],"domain_scores_gemma":[0.5857897,0.25204802,0.106220365,0.013316346,0.03676461,0.0058609024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012547726,0.00027150643,0.00038426148,0.0056483885,0.0015313048,0.0015217441,0.0014100644,0.00098695,0.0009938362],"category_scores_gemma":[0.171848,0.00044549373,0.00029359612,0.0037264163,0.002299875,0.0030420248,0.0018095779,0.0013854436,0.00032666366],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034369578,0.00073788455,0.8717259,0.00039124262,0.000080526625,0.0015411002,0.07477258,0.0009464118,0.002960462,0.0006336027,0.0013691778,0.04449726],"study_design_scores_gemma":[0.000029489483,0.0008307998,0.95812124,0.00014896244,0.00004269588,0.0012758505,0.029978987,0.0033226067,0.0017240334,0.00028883293,0.0041816416,0.000054987948],"about_ca_topic_score_codex":0.010526067,"about_ca_topic_score_gemma":0.014062537,"teacher_disagreement_score":0.012547726,"about_ca_system_score_codex":0.0024392968,"about_ca_system_score_gemma":0.002380024,"threshold_uncertainty_score":0.06635952},"labels":[],"label_agreement":null},{"id":"W4281253013","doi":"10.1016/j.infsof.2022.106942","title":"What do developers consider magic literals? A smalltalk perspective","year":2022,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; École de Technologie Supérieure","funders":"","keywords":"MAGIC (telescope); Computer science; Source code; Smalltalk; Theoretical computer science; Algorithm; Programming language; Mathematics; Discrete mathematics","score_opus":0.010745282438362223,"score_gpt":0.2522363141021629,"score_spread":0.24149103166380068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281253013","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25048137,0.015384625,0.18643077,0.19331296,0.0009558526,0.00016990228,0.00023488983,0.0009857499,0.35204393],"genre_scores_gemma":[0.94779515,0.0032468296,0.025941867,0.007404017,0.00037378387,0.000091295486,0.0000847392,0.0005768628,0.014485341],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9900169,0.0055937334,0.00038352679,0.00084906514,0.0022094336,0.0009473523],"domain_scores_gemma":[0.92444766,0.058767766,0.0037761037,0.0030981004,0.00726876,0.0026416846],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008411848,0.00053984026,0.00068155356,0.0025291846,0.0034367628,0.0103493715,0.0023852456,0.003739106,0.006774604],"category_scores_gemma":[0.059677307,0.00093956204,0.00034361973,0.002021197,0.00989249,0.021405654,0.0034377116,0.0037866698,0.0011139368],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023214407,0.00016331006,0.012390806,0.0007853078,0.00005566518,0.00087931013,0.06899111,0.0007748293,0.0028207065,0.6681059,0.032415505,0.2123855],"study_design_scores_gemma":[0.00009431217,0.000095385316,0.00589271,0.0011084814,0.00008974124,0.0008876842,0.06866316,0.0025435344,0.0029821105,0.7015548,0.21600747,0.000080751364],"about_ca_topic_score_codex":0.010621118,"about_ca_topic_score_gemma":0.010675959,"teacher_disagreement_score":0.010621118,"about_ca_system_score_codex":0.0026341034,"about_ca_system_score_gemma":0.004293391,"threshold_uncertainty_score":0.044486642},"labels":[],"label_agreement":null},{"id":"W4281398044","doi":"10.1145/3488932.3497769","title":"On Measuring Vulnerable JavaScript Functions in the Wild","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 ACM on Asia Conference on Computer and Communications Security","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"JavaScript; Unobtrusive JavaScript; Computer science; Vulnerability (computing); Secure coding; Exploit; Web application; Popularity; Computer security; Rich Internet application; World Wide Web; Cross-site scripting; Source code; Web application security; Web page; Software security assurance; Programming language; Information security; Web development","score_opus":0.047721530598306035,"score_gpt":0.2654061366848843,"score_spread":0.21768460608657825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281398044","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97449523,0.00080944534,0.016756149,0.00010923719,0.000043203054,0.0001019698,0.0049606673,0.0012150928,0.001508895],"genre_scores_gemma":[0.9578916,0.00039930455,0.027104089,0.00008347582,0.000034807057,0.00013016666,0.013421232,0.000156045,0.00077933335],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.995103,0.0008483627,0.0005385619,0.0014395795,0.001696037,0.00037453818],"domain_scores_gemma":[0.9814449,0.0071036085,0.004900707,0.0026373165,0.0031218058,0.0007915957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025898917,0.0009061512,0.00063838693,0.007467256,0.00063441077,0.0010516273,0.0011019008,0.0012803065,0.00037059022],"category_scores_gemma":[0.015994454,0.0003199897,0.0005773049,0.0040308256,0.000837147,0.0030523662,0.001589747,0.000733919,0.00071162224],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041040158,0.00093184755,0.7644084,0.000803643,0.00028504268,0.00080778036,0.0015651103,0.020671897,0.028105307,0.001685896,0.009577939,0.17074673],"study_design_scores_gemma":[0.000030230474,0.00073233416,0.8240215,0.00017290251,0.00016429636,0.002111144,0.0018263852,0.13042608,0.028688507,0.002394819,0.009303187,0.00012855727],"about_ca_topic_score_codex":0.004722769,"about_ca_topic_score_gemma":0.0063233315,"teacher_disagreement_score":0.007467256,"about_ca_system_score_codex":0.0006944805,"about_ca_system_score_gemma":0.00065057154,"threshold_uncertainty_score":0.013696849},"labels":[],"label_agreement":null},{"id":"W4281617421","doi":"10.1007/s10664-021-10108-z","title":"Evolving software system families in space and time with feature revisions","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Bundesministerium für Digitalisierung und Wirtschaftsstandort; Österreichische Forschungsförderungsgesellschaft; Fundação Carlos Chagas Filho de Amparo à Pesquisa do Estado do Rio de Janeiro; Conselho Nacional de Desenvolvimento Científico e Tecnológico; Austrian Science Fund; Österreichische Nationalstiftung für Forschung, Technologie und Entwicklung","keywords":"Feature (linguistics); Correctness; Software; Computer science; Precision and recall; Feature vector; Software system; Feature model; Data mining; Space (punctuation); Artificial intelligence; Algorithm; Programming language","score_opus":0.00914075863851467,"score_gpt":0.23168815875263116,"score_spread":0.2225474001141165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281617421","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56142074,0.00079241325,0.42325243,0.0002566984,0.000056782286,0.00036736001,0.000774236,0.011048014,0.0020312793],"genre_scores_gemma":[0.63469577,0.00028326648,0.36118582,0.000054361288,0.000023584354,0.00009047897,0.0019257631,0.0005991124,0.0011418973],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9930528,0.0017816155,0.0006627378,0.0014028754,0.002905462,0.00019460723],"domain_scores_gemma":[0.96519667,0.016927266,0.005217644,0.0063546333,0.00574698,0.0005567843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038119457,0.000853244,0.0005946675,0.0043601627,0.00060713745,0.0017951397,0.0009901427,0.0007779566,0.00073396444],"category_scores_gemma":[0.03484239,0.0005949187,0.000818731,0.00284178,0.00065088324,0.0028681417,0.0016223914,0.0008460174,0.0005288979],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060662476,0.00028963556,0.11481662,0.0006803049,0.00020303305,0.0009989283,0.0034148917,0.039642412,0.04765498,0.003194096,0.0029826201,0.7855158],"study_design_scores_gemma":[0.00007993392,0.00077000656,0.08710884,0.00031270244,0.00030904874,0.004517813,0.0017732854,0.7240062,0.13777155,0.011868037,0.031296607,0.00018600834],"about_ca_topic_score_codex":0.0020620045,"about_ca_topic_score_gemma":0.0028306947,"teacher_disagreement_score":0.0043601627,"about_ca_system_score_codex":0.00067963434,"about_ca_system_score_gemma":0.00069395587,"threshold_uncertainty_score":0.020159721},"labels":[],"label_agreement":null},{"id":"W4281627483","doi":"10.1007/s10664-022-10156-z","title":"A qualitative study of developers’ discussions of their problems and joys during the early COVID-19 months","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University; University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; University of Calgary","keywords":"Notice; Preparedness; Documentation; Coronavirus disease 2019 (COVID-19); Loneliness; Work (physics); Software; Psychology; Medical education; Computer science; Public relations; World Wide Web; Engineering; Management; Political science; Medicine; Social psychology","score_opus":0.05312803563303627,"score_gpt":0.33345928796693347,"score_spread":0.2803312523338972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281627483","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98847395,0.00031206643,0.0018182003,0.0018178007,0.00010455578,0.00022078381,0.00030296633,0.0000426601,0.0069070854],"genre_scores_gemma":[0.9925041,0.0002777413,0.0009449667,0.0010604017,0.000038862563,0.0005072454,0.00016600655,0.000074402626,0.0044263136],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9867154,0.008361846,0.0004907317,0.00083437905,0.0018335197,0.0017641281],"domain_scores_gemma":[0.87172925,0.09271057,0.008748489,0.0023069507,0.013740353,0.010764371],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016323544,0.0007262795,0.00083857373,0.0031215185,0.011074431,0.005065698,0.0021400596,0.002968,0.0035147883],"category_scores_gemma":[0.07520336,0.0011172806,0.0003441955,0.0025861901,0.006906516,0.0037853802,0.006858384,0.004808002,0.00074535],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006104075,0.00007769174,0.0064641996,0.00007565801,0.0000029302466,0.00048769565,0.9882472,0.000016800574,0.00081892504,0.00040344172,0.000563664,0.0027808384],"study_design_scores_gemma":[0.000011157686,0.00010583705,0.011407321,0.00016083909,0.0000038490934,0.000139158,0.97934556,0.000057439913,0.00037130498,0.00017197916,0.008200776,0.000024676014],"about_ca_topic_score_codex":0.02379125,"about_ca_topic_score_gemma":0.054063085,"teacher_disagreement_score":0.02379125,"about_ca_system_score_codex":0.006378056,"about_ca_system_score_gemma":0.0066298232,"threshold_uncertainty_score":0.08632821},"labels":[],"label_agreement":null},{"id":"W4281729201","doi":"10.1016/j.infsof.2022.106963","title":"Like, dislike, or just do it? How developers approach software development tasks","year":2022,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software engineering; Computer science; Software development; Software; Human–computer interaction; Systems engineering; Engineering; Programming language","score_opus":0.02146837634590896,"score_gpt":0.24172793939720166,"score_spread":0.2202595630512927,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281729201","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73911685,0.0018851466,0.08554625,0.03306239,0.00037146945,0.0001653002,0.000100489284,0.0008350155,0.1389171],"genre_scores_gemma":[0.9545814,0.00076899596,0.028335873,0.0022639467,0.000043740438,0.00007876365,0.000090941125,0.0004906469,0.013345806],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9811786,0.011473523,0.00054116716,0.0011600232,0.00459535,0.001051329],"domain_scores_gemma":[0.95368594,0.027577858,0.0042635244,0.002548106,0.008325703,0.0035988893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011188171,0.00045519785,0.0003058722,0.0018135959,0.002795808,0.009081741,0.0013942935,0.00319258,0.0030231345],"category_scores_gemma":[0.09320541,0.0008040817,0.00033710268,0.00094702333,0.0041629,0.010607468,0.0028452927,0.002820917,0.0018291152],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019050797,0.0005159237,0.14403474,0.00045388364,0.00014137347,0.0015809013,0.408577,0.0013694594,0.0123916,0.059757475,0.037450336,0.33353677],"study_design_scores_gemma":[0.00015162537,0.0004899959,0.14524007,0.001483451,0.00021132594,0.003270824,0.38609815,0.009003418,0.00656117,0.12657367,0.3205092,0.0004072008],"about_ca_topic_score_codex":0.0076618856,"about_ca_topic_score_gemma":0.014504983,"teacher_disagreement_score":0.011188171,"about_ca_system_score_codex":0.0018665316,"about_ca_system_score_gemma":0.0054556155,"threshold_uncertainty_score":0.059169352},"labels":[],"label_agreement":null},{"id":"W4281811500","doi":"10.1007/s10664-022-10153-2","title":"Works for Me! Cannot Reproduce – A Large Scale Empirical Study of Non-reproducible Bugs","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Dalhousie University","funders":"","keywords":"Software bug; Computer science; Empirical research; Software; Debugging; Eclipse; Software regression; Software engineering; Data science; Software development; Software quality; Programming language; Statistics","score_opus":0.030991629615728825,"score_gpt":0.32073289073010847,"score_spread":0.28974126111437964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281811500","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94520026,0.0016153945,0.014062216,0.015152443,0.00041886457,0.0001827723,0.0013145094,0.0005859067,0.021467648],"genre_scores_gemma":[0.98853195,0.00031438732,0.0031277437,0.002137383,0.00011123024,0.00009122654,0.00066942314,0.00043519217,0.0045815124],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9907906,0.003922503,0.00055563624,0.0014555434,0.0029263783,0.00034935607],"domain_scores_gemma":[0.7262294,0.14398326,0.032816518,0.075453974,0.016830187,0.0046865777],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013277922,0.0005016173,0.0005506806,0.001769731,0.0020028243,0.002547783,0.0018459433,0.0025598893,0.010098525],"category_scores_gemma":[0.18588398,0.00069071294,0.000652021,0.0020760836,0.0044168807,0.0062677665,0.00256245,0.002904385,0.004457865],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010181455,0.0016604488,0.6395112,0.001340607,0.0009228447,0.0015750512,0.020517819,0.0030566845,0.0056141415,0.04858618,0.05541445,0.22078247],"study_design_scores_gemma":[0.0005155413,0.0026547755,0.7123674,0.0016379893,0.0005348653,0.0040498525,0.023436608,0.009988741,0.005254822,0.11931348,0.1198994,0.000346518],"about_ca_topic_score_codex":0.0021744922,"about_ca_topic_score_gemma":0.0028116996,"teacher_disagreement_score":0.98672205,"about_ca_system_score_codex":0.00079553167,"about_ca_system_score_gemma":0.001560508,"threshold_uncertainty_score":0.070221186},"labels":[],"label_agreement":null},{"id":"W4282033849","doi":"10.1145/3542944","title":"Towards Learning Generalizable Code Embeddings Using Task-agnostic Graph Convolutional Networks","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Concordia University","funders":"","keywords":"Computer science; Source code; Downstream (manufacturing); Graph; Benchmarking; Code (set theory); Abstract syntax; Embedding; Task (project management); Artificial intelligence; Syntax; Theoretical computer science; Programming language","score_opus":0.07716174346765048,"score_gpt":0.31820709951536275,"score_spread":0.24104535604771227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4282033849","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.177948,0.0017265733,0.7967585,0.0010004874,0.00023111278,0.000176278,0.0019299305,0.015858592,0.0043705516],"genre_scores_gemma":[0.75645584,0.0009612111,0.22072338,0.00090991886,0.00012710557,0.00027758064,0.011774821,0.0009123827,0.00785773],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99942243,0.00012058632,0.0000264867,0.00026322732,0.000088998604,0.00007830877],"domain_scores_gemma":[0.99861455,0.0005073707,0.0001759083,0.00035916825,0.0002671485,0.000075828844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006479251,0.0024446878,0.00069667335,0.0012819347,0.00033135666,0.00075102324,0.0013798375,0.0012474385,0.0012020733],"category_scores_gemma":[0.0036987402,0.0006007485,0.0011991725,0.0011664835,0.00089131866,0.0028712908,0.0014209768,0.0022308773,0.0010412163],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023243434,0.00034213348,0.0068368977,0.00030226685,0.00018156621,0.00020878052,0.00017882026,0.513875,0.015933335,0.0071939486,0.01601021,0.43870455],"study_design_scores_gemma":[0.000013285043,0.00005679006,0.00053901545,0.000014003113,0.000018921835,0.000038935206,0.000025158895,0.9854665,0.002681197,0.009862641,0.0012731707,0.000010455082],"about_ca_topic_score_codex":0.007650756,"about_ca_topic_score_gemma":0.016184883,"teacher_disagreement_score":0.007650756,"about_ca_system_score_codex":0.0010432334,"about_ca_system_score_gemma":0.0011069787,"threshold_uncertainty_score":0.015212476},"labels":[],"label_agreement":null},{"id":"W4282826452","doi":"10.1109/icse-nier55298.2022.9793511","title":"Better Modeling the Programming World with Code Concept Graphs-augmented Multi-modal Learning","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Identifier; Artificial intelligence; Code (set theory); Process (computing); Domain (mathematical analysis); Machine learning; Programming language","score_opus":0.02634966534807758,"score_gpt":0.2644446414475526,"score_spread":0.238094976099475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4282826452","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054542676,0.00017781364,0.94249797,0.0005199604,0.000024498193,0.00003603354,0.00015718209,0.00094003684,0.0011038462],"genre_scores_gemma":[0.7316679,0.00020250061,0.2654988,0.00025233137,0.0000260589,0.00010216454,0.00046611956,0.00012936762,0.0016547247],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996731,0.00011736035,0.000012931825,0.00010675871,0.000058523387,0.000031227235],"domain_scores_gemma":[0.9985207,0.00086680066,0.00014223639,0.00022486839,0.00017948583,0.00006597691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007264297,0.0006253858,0.00043878533,0.0008218729,0.00028149437,0.0010992492,0.0013931103,0.00095819915,0.0015598536],"category_scores_gemma":[0.003590442,0.00035400141,0.0009495079,0.00071191014,0.0008238821,0.0035865954,0.0012131458,0.0024351557,0.00029111127],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007468506,0.00013263384,0.0019678187,0.00008388062,0.000043692053,0.000075416676,0.00016463957,0.89681506,0.0036287836,0.02425496,0.0012310871,0.07152733],"study_design_scores_gemma":[0.000001827672,0.000009892852,0.0000753871,0.0000026721311,0.0000023556297,0.0000066190773,0.000008224889,0.98910564,0.00031230308,0.010279206,0.00019305869,0.0000026987655],"about_ca_topic_score_codex":0.0069229836,"about_ca_topic_score_gemma":0.011365614,"teacher_disagreement_score":0.0069229836,"about_ca_system_score_codex":0.000947382,"about_ca_system_score_gemma":0.00085168367,"threshold_uncertainty_score":0.013765395},"labels":[],"label_agreement":null},{"id":"W4282834229","doi":"10.2174/2666255816666220609110712","title":"Systematic Review of Machine Learning-Based Open-Source SoftwareMaintenance Effort Estimation","year":2022,"lang":"en","type":"article","venue":"Recent Advances in Computer Science and Communications","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Machine learning; Data mining; Decision tree; Software; Context (archaeology); Support vector machine; Artificial intelligence; Empirical research; Feature selection; Statistics","score_opus":0.019911980617795505,"score_gpt":0.31668119498459185,"score_spread":0.2967692143667964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4282834229","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015502173,0.99469477,0.0012585953,0.00054789917,0.00016429923,0.00050632947,0.00078949565,0.00002584611,0.0004625336],"genre_scores_gemma":[0.021983791,0.96915674,0.005234364,0.0007335142,0.00015553595,0.0014690143,0.001072356,0.000023922077,0.00017077607],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.96992075,0.009556314,0.011827129,0.002140833,0.006164126,0.0003908012],"domain_scores_gemma":[0.80197746,0.15081252,0.025099073,0.0032755735,0.017814819,0.0010206144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028670006,0.0020849193,0.00743571,0.033660773,0.0008981075,0.0033588237,0.003737602,0.0023859912,0.003229344],"category_scores_gemma":[0.15637732,0.0013793986,0.0075562396,0.019983832,0.0015159833,0.005039962,0.002636951,0.0017230781,0.0005003716],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012126106,0.00003872267,0.0017885317,0.86982733,0.0050374465,0.00013876501,0.00038575788,0.0003977696,0.00021228645,0.0005838635,0.0021964726,0.119271874],"study_design_scores_gemma":[0.00009285099,0.00025451908,0.0048905476,0.9444461,0.023422418,0.00032304312,0.00040560964,0.00032522052,0.00037513385,0.0007122552,0.024696032,0.000056323952],"about_ca_topic_score_codex":0.00862702,"about_ca_topic_score_gemma":0.022870349,"teacher_disagreement_score":0.033660773,"about_ca_system_score_codex":0.0054049026,"about_ca_system_score_gemma":0.025509506,"threshold_uncertainty_score":0.15162325},"labels":[],"label_agreement":null},{"id":"W4282838525","doi":"10.1109/icse-nier55298.2022.9793515","title":"Supporting program comprehension by generating abstract code summary tree","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Program comprehension; Programming language; Tree (set theory); Code (set theory); Software; Mathematics; Software system","score_opus":0.025275417511071867,"score_gpt":0.31248482074930567,"score_spread":0.2872094032382338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4282838525","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044770587,0.00027131202,0.88566977,0.00040562698,0.00009261701,0.000668578,0.0056139496,0.05952316,0.00298438],"genre_scores_gemma":[0.115970425,0.000265094,0.8571924,0.00017923646,0.000053012984,0.0004889258,0.018222429,0.004540456,0.003088022],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99869746,0.0002654156,0.00014002033,0.00030915975,0.0005225553,0.000065482236],"domain_scores_gemma":[0.99016094,0.0046472806,0.0009731133,0.0012763549,0.0027172768,0.00022501282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012811322,0.0013954983,0.0006943842,0.0034531595,0.0005748815,0.0013942849,0.001224852,0.00096857647,0.0050476836],"category_scores_gemma":[0.012653107,0.00055779424,0.0008970525,0.0022951958,0.00038710985,0.002930671,0.0011984634,0.0010906693,0.0025029927],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066129153,0.00043508472,0.0066439505,0.0018874058,0.0000855074,0.0012082541,0.0052774325,0.021210957,0.09148811,0.009605021,0.053487074,0.8080099],"study_design_scores_gemma":[0.00026050868,0.0007796683,0.009287909,0.00035319332,0.0002929172,0.001371811,0.0020304713,0.6262927,0.13641332,0.03676607,0.18591529,0.00023606862],"about_ca_topic_score_codex":0.0026579038,"about_ca_topic_score_gemma":0.0042540277,"teacher_disagreement_score":0.0050476836,"about_ca_system_score_codex":0.0005317387,"about_ca_system_score_gemma":0.0014569897,"threshold_uncertainty_score":0.016886175},"labels":[],"label_agreement":null},{"id":"W4283077896","doi":"10.1016/j.infsof.2022.106985","title":"A three-stage transfer learning framework for multi-source cross-project software defect prediction","year":2022,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Multi-source; Computer science; Transfer of learning; Source code; Weighting; Data source; Open source; Selection (genetic algorithm); Merge (version control); Data mining; Open source software; Field (mathematics); Water source; Machine learning; Software; Artificial intelligence; Information retrieval; Statistics","score_opus":0.025754898635890856,"score_gpt":0.2949621594112539,"score_spread":0.26920726077536306,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283077896","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015360572,0.00023699379,0.9818611,0.00022704086,0.000035968398,0.0001706841,0.00016372754,0.0013515602,0.00059231324],"genre_scores_gemma":[0.5881259,0.0003574418,0.40449342,0.00030046905,0.00013870955,0.0011452525,0.0014466954,0.00022672917,0.003765441],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99766433,0.000829299,0.00015333133,0.00066314096,0.00046129277,0.00022861143],"domain_scores_gemma":[0.99578,0.0019793129,0.00032320977,0.00043613208,0.0012759322,0.00020530066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062061427,0.0016351996,0.0014266736,0.0028642088,0.0007882741,0.0012136301,0.0042802426,0.0019863264,0.0021468266],"category_scores_gemma":[0.00827405,0.0007622797,0.0017407591,0.002475788,0.0009318796,0.0035331121,0.0028497363,0.0026625367,0.0008421034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021326971,0.00048920384,0.0061006173,0.00016801668,0.00024321629,0.00019378164,0.00027003832,0.5842424,0.0023006785,0.0058202758,0.003843456,0.396115],"study_design_scores_gemma":[0.0000074329564,0.000040353345,0.0003020702,0.000004747647,0.000009549978,0.000011876314,0.000013766087,0.9959264,0.00038939485,0.002999478,0.0002865226,0.000008420595],"about_ca_topic_score_codex":0.010285189,"about_ca_topic_score_gemma":0.0080303615,"teacher_disagreement_score":0.010285189,"about_ca_system_score_codex":0.0014743066,"about_ca_system_score_gemma":0.0020084535,"threshold_uncertainty_score":0.032821596},"labels":[],"label_agreement":null},{"id":"W4283653799","doi":"10.5281/zenodo.6760334","title":"Impact of design choices on the quality of multi-modal transformer based embedding for bug localization","year":2023,"lang":"en","type":"paratext","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Modal; Transformer; Embedding; Computer science; Quality (philosophy); Electronic engineering; Reliability engineering; Engineering; Electrical engineering; Artificial intelligence; Voltage; Materials science; Physics","score_opus":0.1213326945007197,"score_gpt":0.3581642551044042,"score_spread":0.2368315606036845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283653799","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64410454,0.016763544,0.066477194,0.0026479284,0.0011745495,0.00031579353,0.19222815,0.057623804,0.018664496],"genre_scores_gemma":[0.5743105,0.0017293867,0.06563323,0.00065557635,0.00010193703,0.000225916,0.34808558,0.0024211085,0.006836828],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966537,0.0009807542,0.00024728698,0.000797673,0.0011372673,0.00018337103],"domain_scores_gemma":[0.9913796,0.004499273,0.0006265263,0.0021515358,0.0011419345,0.00020109513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002671401,0.0019013373,0.00082468404,0.0022575916,0.00039204926,0.0013159317,0.0010617281,0.0013530543,0.0033293068],"category_scores_gemma":[0.01602471,0.00033162086,0.00143768,0.001668343,0.00070944626,0.0019346441,0.0013708201,0.0010695339,0.0023185087],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034025125,0.00086648425,0.057174917,0.007932269,0.0013695508,0.0008466184,0.00034682793,0.15075006,0.02725525,0.004902539,0.3342993,0.41085374],"study_design_scores_gemma":[0.0012720478,0.002557298,0.06590118,0.0011431955,0.0016676657,0.002593034,0.0009203995,0.6024796,0.07292041,0.022537697,0.22573051,0.0002770615],"about_ca_topic_score_codex":0.0031131478,"about_ca_topic_score_gemma":0.007936451,"teacher_disagreement_score":0.0033293068,"about_ca_system_score_codex":0.0007378236,"about_ca_system_score_gemma":0.00073525537,"threshold_uncertainty_score":0.01412791},"labels":[],"label_agreement":null},{"id":"W4283737447","doi":"10.47839/ijc.21.2.2587","title":"Software Reusability Estimation based on Dynamic Metrics using Soft Computing Techniques","year":2022,"lang":"en","type":"article","venue":"International Journal of Computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Reusability; Computer science; Soft computing; Mean squared error; Fuzzy logic; Software; Data mining; Artificial neural network; Object-oriented programming; Artificial intelligence; Machine learning; Programming language; Statistics; Mathematics","score_opus":0.022148613568665373,"score_gpt":0.33307436842799204,"score_spread":0.31092575485932666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283737447","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21376246,0.0004529916,0.78313285,0.000122034224,0.00003023311,0.0000674791,0.000114071474,0.00054360233,0.0017743238],"genre_scores_gemma":[0.93335795,0.00015150006,0.065676205,0.000014880312,0.000011954045,0.000057752786,0.0001436312,0.000023765659,0.00056234753],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980203,0.00038119918,0.000207803,0.00032747292,0.0009766025,0.00008654483],"domain_scores_gemma":[0.99525046,0.0019039126,0.0011110812,0.00042239582,0.0012114078,0.00010067361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017175644,0.0010219758,0.0007863767,0.0061264527,0.0003078349,0.0011903403,0.0007358417,0.0006539103,0.00042332063],"category_scores_gemma":[0.00898469,0.0003420837,0.0008032408,0.0027383598,0.000492559,0.0016350236,0.0007244739,0.0005934358,0.00015950989],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000115247996,0.00020540829,0.045459367,0.00018448131,0.00030217864,0.00018915214,0.00021499081,0.632112,0.018673258,0.0052209073,0.00045089543,0.29687208],"study_design_scores_gemma":[0.0000023231948,0.000047164332,0.006128026,0.00001477105,0.00001864595,0.00003932813,0.000023046641,0.9888621,0.0026700853,0.0019633675,0.00021592881,0.000015296415],"about_ca_topic_score_codex":0.005165315,"about_ca_topic_score_gemma":0.0047394405,"teacher_disagreement_score":0.0061264527,"about_ca_system_score_codex":0.00095502555,"about_ca_system_score_gemma":0.0006053835,"threshold_uncertainty_score":0.010270476},"labels":[],"label_agreement":null},{"id":"W4283796458","doi":"10.1609/aaai.v36i11.21625","title":"Code Representation Learning Using Prüfer Sequences (Student Abstract)","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Representation (politics); Sequence (biology); Syntax; Abstract syntax; Code (set theory); Encoding (memory); Natural language processing; Source code; Artificial intelligence; Lossless compression; Abstract syntax tree; Programming language; Tree (set theory); Scheme (mathematics); Artificial neural network; Theoretical computer science; Data compression; Mathematics","score_opus":0.14454852823829395,"score_gpt":0.36782651393995397,"score_spread":0.22327798570166002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283796458","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0834218,0.000171873,0.9090769,0.00043717405,0.000088901535,0.0000883588,0.00036749378,0.004048568,0.002298927],"genre_scores_gemma":[0.5737356,0.0002160309,0.41907513,0.00019791297,0.000048786842,0.00017765506,0.0013531788,0.0002441716,0.004951478],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996568,0.00009230705,0.000020123283,0.00010484647,0.00009458413,0.00003134503],"domain_scores_gemma":[0.9987809,0.00045995356,0.00015702899,0.0002858466,0.0002557183,0.000060459406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005167469,0.0004806437,0.0002895507,0.00060117204,0.00024704673,0.00055040774,0.0007593586,0.0005845853,0.0031535008],"category_scores_gemma":[0.004273302,0.00018751221,0.00036429716,0.0005598912,0.00047119617,0.0023151075,0.00080906134,0.0011003378,0.0010400113],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044869597,0.00019793732,0.0020970257,0.00016049216,0.00003054173,0.00017533683,0.00021517285,0.17565098,0.032179274,0.03898038,0.0072945375,0.74256957],"study_design_scores_gemma":[0.000017768656,0.00014164146,0.0002484951,0.000014632792,0.000008907357,0.000077047545,0.000020402817,0.9599905,0.016495964,0.019700151,0.0032719157,0.00001263829],"about_ca_topic_score_codex":0.0028395609,"about_ca_topic_score_gemma":0.0025244758,"teacher_disagreement_score":0.0031535008,"about_ca_system_score_codex":0.00046416634,"about_ca_system_score_gemma":0.0007757916,"threshold_uncertainty_score":0.010549545},"labels":[],"label_agreement":null},{"id":"W4284670329","doi":"10.1145/3510003.3510076","title":"Social science theories in software engineering research","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 44th International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Social software engineering; Software; Software Engineering Process Group; Computer science; Software development; Management science; Engineering ethics; Field (mathematics); Epistemology; Software development process; Software construction; Engineering; Mathematics","score_opus":0.047533180329221374,"score_gpt":0.3234322068789773,"score_spread":0.2758990265497559,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284670329","genre_codex":"other","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07895622,0.09571316,0.2613403,0.14211728,0.0022008498,0.0007022428,0.0004267425,0.00034432893,0.4181988],"genre_scores_gemma":[0.86423606,0.058785234,0.051127847,0.009675725,0.0026270677,0.0007981271,0.00029407887,0.00014306177,0.01231286],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9781953,0.013416261,0.00094241684,0.0014956932,0.005168413,0.0007818776],"domain_scores_gemma":[0.8602081,0.12373627,0.0056718285,0.0030158828,0.005549036,0.0018188863],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.025313722,0.0011511463,0.0012602552,0.015753567,0.0041864817,0.0099148955,0.0019300596,0.0043211333,0.0064039635],"category_scores_gemma":[0.04176472,0.00065945025,0.0012515481,0.014564938,0.021211484,0.013932861,0.0050655324,0.0054000807,0.0011981792],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011135649,0.00007831446,0.0039016139,0.00049522746,0.00007402335,0.00018916537,0.005809262,0.0027057116,0.00007502395,0.94862515,0.0036626575,0.034372702],"study_design_scores_gemma":[0.000016001588,0.000030085097,0.0019339854,0.00057204685,0.000022103786,0.00010441533,0.003082539,0.0023726427,0.00011855414,0.9495999,0.042121895,0.000025876572],"about_ca_topic_score_codex":0.0041887867,"about_ca_topic_score_gemma":0.0033574002,"teacher_disagreement_score":0.9958135,"about_ca_system_score_codex":0.0074325437,"about_ca_system_score_gemma":0.006741716,"threshold_uncertainty_score":0.13387334},"labels":[],"label_agreement":null},{"id":"W4284670745","doi":"10.1145/3510003.3510115","title":"Inferring and applying type changes","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 44th International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada; University of Colorado Boulder; National Science Foundation","keywords":"Code refactoring; Computer science; Plug-in; Programming language; Maintainability; Type (biology); Precision and recall; Data type; Software engineering; Information retrieval; Data mining; Artificial intelligence; Software","score_opus":0.03354994806902248,"score_gpt":0.2610165799266521,"score_spread":0.22746663185762964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284670745","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11875535,0.0017718216,0.77815324,0.0013564508,0.00089068786,0.0010438057,0.011548524,0.0778802,0.008599901],"genre_scores_gemma":[0.25200614,0.0014139264,0.7063318,0.0010847972,0.00029912114,0.0004163557,0.023822686,0.00884079,0.0057845167],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.980243,0.003322518,0.0018975331,0.005345596,0.00833675,0.0008546785],"domain_scores_gemma":[0.94936156,0.025407739,0.003925497,0.0116526205,0.009105073,0.0005475162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008309321,0.0021833614,0.001580775,0.0071194125,0.0014245707,0.0037277017,0.0030834766,0.0021674216,0.0027903002],"category_scores_gemma":[0.07982732,0.0017780826,0.0027453743,0.0030610221,0.0013624883,0.003858324,0.003251539,0.003722255,0.0025566234],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046119982,0.0003979106,0.12168821,0.002044215,0.0005081052,0.002383844,0.0029288104,0.024560085,0.031875174,0.011951217,0.043692213,0.75750905],"study_design_scores_gemma":[0.0002049811,0.00038757533,0.047225583,0.0011317536,0.0011447427,0.0041060154,0.0019489543,0.52663296,0.14563434,0.044941653,0.22614527,0.0004962073],"about_ca_topic_score_codex":0.010089878,"about_ca_topic_score_gemma":0.016876254,"teacher_disagreement_score":0.010089878,"about_ca_system_score_codex":0.0013839047,"about_ca_system_score_gemma":0.0044272407,"threshold_uncertainty_score":0.04394442},"labels":[],"label_agreement":null},{"id":"W4284689801","doi":"10.1145/3510003.3510157","title":"Automated handling of anaphoric ambiguity in requirements","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 44th International Conference on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds National de la Recherche Luxembourg","keywords":"Ambiguity; Computer science; Anaphora (linguistics); Artificial intelligence; Natural language processing; Coreference; Natural language; Natural language understanding; Interpretation (philosophy); Ambiguity resolution; Automation; Machine translation; Feature (linguistics); Resolution (logic); Machine learning; Programming language; Linguistics","score_opus":0.0372657344313989,"score_gpt":0.28603588832004295,"score_spread":0.24877015388864404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284689801","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064432904,0.0004205688,0.91422963,0.0012781761,0.000039967996,0.00046788037,0.000686894,0.013083723,0.0053602555],"genre_scores_gemma":[0.3045946,0.00023179418,0.6903859,0.0003677434,0.000033610348,0.00019048313,0.0017696955,0.0006348733,0.001791301],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98780257,0.0057305316,0.0007252051,0.0014240202,0.0039877445,0.00032984573],"domain_scores_gemma":[0.97518456,0.014819437,0.0025468483,0.004091058,0.0031171916,0.00024096349],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005489571,0.0014268924,0.0008518225,0.0027365126,0.0010108805,0.0022479221,0.0016253161,0.0015284751,0.002427839],"category_scores_gemma":[0.032909725,0.0009327004,0.0015558329,0.0016402362,0.0010375951,0.0043901997,0.0031276585,0.0023574077,0.0014955065],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024319465,0.00041324602,0.0071193767,0.0012865567,0.0001578377,0.0009468944,0.0027216398,0.08308984,0.04058652,0.023101911,0.01537163,0.82496136],"study_design_scores_gemma":[0.000079320016,0.00018641725,0.0049907663,0.00020146114,0.00011507716,0.001671576,0.0016470724,0.82471997,0.05791468,0.06565583,0.042692646,0.00012525408],"about_ca_topic_score_codex":0.002268629,"about_ca_topic_score_gemma":0.003828986,"teacher_disagreement_score":0.005489571,"about_ca_system_score_codex":0.0010299924,"about_ca_system_score_gemma":0.0030125969,"threshold_uncertainty_score":0.029031992},"labels":[],"label_agreement":null},{"id":"W4284882237","doi":"10.1007/s10664-022-10119-4","title":"FIXME: synchronize with database! An empirical study of data access self-admitted technical debt","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Technical debt; Code refactoring; Computer science; Commit; Maintainability; Database; Data access; Software development; Software; Data science; Software engineering; Operating system","score_opus":0.06157596337977576,"score_gpt":0.35714592806012,"score_spread":0.29556996468034424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284882237","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99648595,0.00009784149,0.0006240519,0.0004572428,0.000009590053,0.000017964107,0.00025080692,0.00007136282,0.001985087],"genre_scores_gemma":[0.9982602,0.00004582706,0.00039132594,0.00010457865,0.000013277841,0.000013426953,0.00030207002,0.00003849149,0.0008308],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99709034,0.0015764888,0.00022589105,0.0002964195,0.0005733564,0.00023747118],"domain_scores_gemma":[0.85123694,0.106118105,0.023534797,0.011154841,0.0039941724,0.003961211],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060131955,0.00022423542,0.00026957784,0.0013265477,0.0011179487,0.002298464,0.0012023099,0.0015309979,0.0048948172],"category_scores_gemma":[0.09435685,0.00041081605,0.00018526743,0.0019960941,0.0016679452,0.0056459066,0.0021968987,0.002964011,0.000994911],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094149954,0.0021947755,0.92972374,0.00015411396,0.0000863527,0.0006224796,0.010702951,0.0014661413,0.0012897394,0.01252768,0.0073598027,0.032930795],"study_design_scores_gemma":[0.00028321156,0.0020269917,0.9014994,0.0002347071,0.00014911253,0.0016958785,0.03178142,0.023467032,0.0038039563,0.01227859,0.02265626,0.00012345357],"about_ca_topic_score_codex":0.0052152895,"about_ca_topic_score_gemma":0.005116061,"teacher_disagreement_score":0.0060131955,"about_ca_system_score_codex":0.00077181193,"about_ca_system_score_gemma":0.0008004588,"threshold_uncertainty_score":0.031801224},"labels":[],"label_agreement":null},{"id":"W4285229229","doi":"10.1007/978-3-031-04829-6_18","title":"Towards a Taxonomy of Software Maintainability Predictors: A Detailed View","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Maintainability; Terminology; Ambiguity; Taxonomy (biology); Computer science; Software quality; Software; Empirical research; Software engineering; Data science; Software development; Linguistics; Mathematics; Statistics; Programming language","score_opus":0.019554233127489964,"score_gpt":0.230527297642859,"score_spread":0.21097306451536904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285229229","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023125485,0.024126016,0.9189713,0.0055242935,0.00022307638,0.00022167669,0.0020830666,0.0019090754,0.023815932],"genre_scores_gemma":[0.18830577,0.02816771,0.76574975,0.0010692297,0.00068307103,0.0003908305,0.0043958835,0.00041157563,0.010826194],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99811614,0.00033079027,0.0002683741,0.0003974284,0.0007403407,0.00014696027],"domain_scores_gemma":[0.99296534,0.0029667611,0.00061467645,0.00079448876,0.0023628972,0.00029582588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027223988,0.0014527818,0.0013860798,0.009975165,0.0014344847,0.00901771,0.0026509522,0.0019492389,0.0036307664],"category_scores_gemma":[0.0088772895,0.0010281609,0.0013824939,0.00977905,0.0018159826,0.0147847,0.0019807203,0.0033977828,0.0021868588],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010258631,0.00028201338,0.01772263,0.00090381014,0.00013173034,0.00036260963,0.00095213746,0.015345683,0.005220389,0.4790266,0.015932854,0.46401703],"study_design_scores_gemma":[0.000022011905,0.00035392045,0.010359834,0.0010004524,0.00022864708,0.0015041138,0.0010779074,0.14725322,0.0036958996,0.74695563,0.08736316,0.00018512856],"about_ca_topic_score_codex":0.0038102707,"about_ca_topic_score_gemma":0.0034367645,"teacher_disagreement_score":0.009975165,"about_ca_system_score_codex":0.0015286468,"about_ca_system_score_gemma":0.0018855832,"threshold_uncertainty_score":0.014397562},"labels":[],"label_agreement":null},{"id":"W4285240847","doi":"10.1007/978-3-031-08760-8_45","title":"Digging Deeper into the State of the Practice for Domain Specific Research Software","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Documentation; Usability; Software; Best practice; Domain (mathematical analysis); Software engineering; Ranking (information retrieval); Data science; World Wide Web; Information retrieval; Human–computer interaction; Programming language","score_opus":0.03214978094487921,"score_gpt":0.31231830375869574,"score_spread":0.28016852281381655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285240847","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0095214555,0.08823263,0.3628689,0.20623098,0.0024818368,0.00012358377,0.00022824049,0.001466854,0.32884553],"genre_scores_gemma":[0.29800797,0.119994536,0.357737,0.036301512,0.0033652876,0.00040111202,0.00077234494,0.0023853828,0.18103483],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9913936,0.004681502,0.00046674904,0.00083926215,0.0021937175,0.00042510167],"domain_scores_gemma":[0.9707818,0.01927893,0.0006103547,0.00652148,0.0020422107,0.00076518406],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012697415,0.00078541925,0.00084109243,0.0034039048,0.0019503587,0.01624431,0.002386031,0.0057440316,0.015024696],"category_scores_gemma":[0.019187136,0.00086486817,0.0008344118,0.004384886,0.01813247,0.036643222,0.005862881,0.011442564,0.0047991346],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009367096,0.000019137617,0.00008042367,0.00019445045,0.0000055582846,0.000021016274,0.0018898898,0.00023826954,0.00031890767,0.92697406,0.009261561,0.06098745],"study_design_scores_gemma":[0.000006814708,0.000023098974,0.00018045816,0.00099948,0.0000065373542,0.00011499195,0.0010656825,0.0014152536,0.00046073875,0.67288345,0.32282525,0.000018145694],"about_ca_topic_score_codex":0.0033093889,"about_ca_topic_score_gemma":0.004507261,"teacher_disagreement_score":0.9873026,"about_ca_system_score_codex":0.004310057,"about_ca_system_score_gemma":0.006513002,"threshold_uncertainty_score":0.06715113},"labels":[],"label_agreement":null},{"id":"W4285310629","doi":"10.1007/978-3-031-08129-3_1","title":"Fine-Grained Analysis of Similar Code Snippets","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Code (set theory); Information retrieval; World Wide Web; Programming language","score_opus":0.019778436055914693,"score_gpt":0.26693290289161986,"score_spread":0.24715446683570516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285310629","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60868275,0.0018952471,0.33041945,0.0004885714,0.0004113252,0.00068140234,0.010555155,0.030414347,0.016451733],"genre_scores_gemma":[0.69520926,0.00062465,0.2639819,0.00034666999,0.00015527953,0.00030119528,0.017038442,0.0052895583,0.017053042],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987746,0.00007411085,0.00008465808,0.00035161973,0.0005695387,0.00014556883],"domain_scores_gemma":[0.99619895,0.0014307545,0.0004599989,0.0007605372,0.0009894569,0.00016034735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004409483,0.0006860204,0.0006836271,0.004475993,0.0010089845,0.0014672916,0.0010994944,0.0010044196,0.0073588802],"category_scores_gemma":[0.004441003,0.00040534016,0.0012427185,0.0031246005,0.0006410225,0.001588604,0.0015268102,0.0009112274,0.0030633016],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015601914,0.000499062,0.04089282,0.0013529412,0.00038072476,0.003428255,0.0013179327,0.02590303,0.30437684,0.013887621,0.022096688,0.5843038],"study_design_scores_gemma":[0.00012949173,0.00085644255,0.10452191,0.00038856248,0.0006336013,0.004984083,0.001717483,0.5532917,0.22687912,0.042085525,0.06426058,0.00025153908],"about_ca_topic_score_codex":0.0045218095,"about_ca_topic_score_gemma":0.009947608,"teacher_disagreement_score":0.0073588802,"about_ca_system_score_codex":0.00062870403,"about_ca_system_score_gemma":0.0010321214,"threshold_uncertainty_score":0.02461791},"labels":[],"label_agreement":null},{"id":"W4285394558","doi":"10.1145/3546945","title":"Assessing the Alignment between the Information Needs of Developers and the Documentation of Programming Languages: A Case Study on Rust","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Huawei Technologies (Canada)","funders":"","keywords":"Documentation; Computer science; Technical documentation; Application programming interface; World Wide Web; Internal documentation; Set (abstract data type); Baseline (sea); Software engineering; Information retrieval; Programming language; Software development; Software","score_opus":0.0778948728389709,"score_gpt":0.3699735239963255,"score_spread":0.2920786511573546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285394558","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9933356,0.00020043358,0.00465961,0.00023015244,0.000009150506,0.00010000279,0.00036140668,0.00009770216,0.0010059358],"genre_scores_gemma":[0.9797323,0.00018608687,0.017196352,0.00010595369,0.000024998764,0.00016205621,0.0016497283,0.00008695422,0.00085549604],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9932881,0.004024745,0.00048730045,0.00094117346,0.0009986331,0.00026003813],"domain_scores_gemma":[0.91566086,0.06723779,0.0062963655,0.003176809,0.0060944343,0.0015337202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0083435,0.00036065897,0.00060069206,0.003043409,0.0011932311,0.0014867116,0.00072943035,0.0011273103,0.0006952808],"category_scores_gemma":[0.041403756,0.00030461125,0.00046338805,0.0038491474,0.0008770524,0.0024436074,0.0015817342,0.001471685,0.00034133586],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016319362,0.0036469891,0.5815503,0.0019433225,0.0002971357,0.0062808534,0.09014054,0.015105577,0.0215234,0.0032133826,0.009938231,0.26472828],"study_design_scores_gemma":[0.00032473635,0.0030150865,0.7171318,0.00053263205,0.0003224562,0.0044911588,0.07279516,0.13957971,0.024449207,0.005724139,0.031375613,0.0002583308],"about_ca_topic_score_codex":0.0068254224,"about_ca_topic_score_gemma":0.012328102,"teacher_disagreement_score":0.0083435,"about_ca_system_score_codex":0.0013886563,"about_ca_system_score_gemma":0.0009960762,"threshold_uncertainty_score":0.0441252},"labels":[],"label_agreement":null},{"id":"W4285397054","doi":"10.48550/arxiv.2207.05132","title":"Dev2vec: Representing Domain Expertise of Developers in an Embedding Space","year":2022,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Embedding; Computer science; Domain (mathematical analysis); Space (punctuation); Software; Data science; Exploit; Software engineering; World Wide Web; Information retrieval; Artificial intelligence; Computer security; Programming language; Mathematics","score_opus":0.08073216102822943,"score_gpt":0.23874237143958046,"score_spread":0.15801021041135105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285397054","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3006064,0.0038195278,0.6607438,0.0015339758,0.0010984609,0.00034276853,0.016934697,0.009409578,0.005510857],"genre_scores_gemma":[0.792918,0.0014507648,0.16068882,0.00044629222,0.00041877126,0.00046723304,0.036277942,0.00059023104,0.0067420034],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987832,0.00037534334,0.00009770634,0.00036639246,0.00024603013,0.00013132987],"domain_scores_gemma":[0.9974336,0.0013581642,0.00020277052,0.00031093985,0.0005787634,0.000115695366],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010966737,0.0024926185,0.00075423927,0.0029657404,0.00041422236,0.0012215502,0.00082386367,0.001111314,0.002206369],"category_scores_gemma":[0.0046819956,0.00034783126,0.0007580265,0.0023550582,0.0006926034,0.001818323,0.0011886169,0.0017686851,0.0015564464],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081415306,0.0006678882,0.024795149,0.0010203084,0.00034937964,0.00055934314,0.0009761519,0.09570122,0.020558672,0.0069574444,0.08435275,0.76324755],"study_design_scores_gemma":[0.000050106584,0.00022994184,0.0066740005,0.00009042429,0.00006923722,0.0003038082,0.0002672913,0.96185714,0.008813921,0.008418049,0.013164233,0.00006189834],"about_ca_topic_score_codex":0.0057187523,"about_ca_topic_score_gemma":0.008693486,"teacher_disagreement_score":0.0057187523,"about_ca_system_score_codex":0.0006956712,"about_ca_system_score_gemma":0.00086035166,"threshold_uncertainty_score":0.011370897},"labels":[],"label_agreement":null},{"id":"W4285464731","doi":"10.32920/ryerson.14654229.v1","title":"Analysis on Relationship between Code Quality and Code Coverage in an XP Environment: a Case Study on the SWURV Project","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Software quality; Code (set theory); Quality (philosophy); Software; Extreme programming; Residual; Code coverage; Code smell; Development testing; Process (computing); Static program analysis; Source code; Root cause; Reliability engineering; Software engineering; Software development; Software development process; Programming language; Engineering; Algorithm; Set (abstract data type)","score_opus":0.2052220897337871,"score_gpt":0.4005590867465069,"score_spread":0.1953369970127198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285464731","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99833894,0.000039642076,0.0011917925,0.000052873216,7.5903984e-7,0.000019939751,0.000032985386,0.00000884239,0.0003142903],"genre_scores_gemma":[0.99682957,0.00005793255,0.0027001973,0.000010268084,0.0000019025318,0.000021542883,0.00011872838,0.000009555704,0.00025029143],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9939824,0.0030776132,0.0002939432,0.0006868734,0.0016025148,0.00035671788],"domain_scores_gemma":[0.9064311,0.08040418,0.00478001,0.0016908474,0.0057941745,0.0008997282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005938465,0.00033259668,0.00037885198,0.002615379,0.0007270839,0.001013446,0.00078839064,0.00089799234,0.00063778757],"category_scores_gemma":[0.026331235,0.00027597471,0.0005903758,0.0019535015,0.0010157595,0.0009752408,0.0007838506,0.00095591036,0.000099291436],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031951966,0.0017855961,0.90095896,0.00029723404,0.00018519507,0.007065121,0.012835668,0.007832821,0.0042700614,0.0016744524,0.0006604385,0.062114883],"study_design_scores_gemma":[0.000054600623,0.0021139246,0.9074095,0.00012224252,0.00015752605,0.0034369435,0.014623149,0.063835,0.0054980507,0.0010851339,0.0016061567,0.000057803885],"about_ca_topic_score_codex":0.0070783356,"about_ca_topic_score_gemma":0.007991226,"teacher_disagreement_score":0.0070783356,"about_ca_system_score_codex":0.0010370486,"about_ca_system_score_gemma":0.0008317623,"threshold_uncertainty_score":0.031405985},"labels":[],"label_agreement":null},{"id":"W4285490440","doi":"10.1145/3533767.3534220","title":"DocTer: documentation-guided fuzzing for testing deep learning API functions","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Science Foundation","keywords":"Fuzz testing; Computer science; Documentation; Parsing; Function (biology); Constraint (computer-aided design); Software bug; Artificial intelligence; Dependency (UML); Software; Programming language; Machine learning; Data mining","score_opus":0.060390053064120076,"score_gpt":0.33188639791175195,"score_spread":0.27149634484763185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285490440","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051333994,0.00052684045,0.76540345,0.0008404351,0.0002504396,0.0004085841,0.005575373,0.1676906,0.007970275],"genre_scores_gemma":[0.37607998,0.0003000507,0.58701926,0.00082962576,0.00007397209,0.0006709633,0.010043328,0.018766347,0.0062164664],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.992991,0.002302941,0.00070219114,0.0012400814,0.0022608363,0.00050297275],"domain_scores_gemma":[0.9694574,0.019616313,0.0012506293,0.00691068,0.0023680967,0.00039681725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048174094,0.0023602792,0.0007966589,0.0022932705,0.0007586462,0.0020858522,0.0039539235,0.0020895298,0.018619152],"category_scores_gemma":[0.04420583,0.0014928223,0.0016325488,0.0010882082,0.0021074233,0.004932171,0.003255012,0.0026787745,0.0043544983],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020004252,0.00087328575,0.019638512,0.0024550776,0.00033718135,0.0014575919,0.0008702662,0.12623578,0.048643027,0.049896024,0.1292096,0.61838317],"study_design_scores_gemma":[0.00042526124,0.00030765645,0.0025357802,0.00033610244,0.000091423084,0.0007046236,0.00019790202,0.7960191,0.102413245,0.058686696,0.03815477,0.0001274242],"about_ca_topic_score_codex":0.0036605992,"about_ca_topic_score_gemma":0.008108934,"teacher_disagreement_score":0.018619152,"about_ca_system_score_codex":0.0015400068,"about_ca_system_score_gemma":0.0028008444,"threshold_uncertainty_score":0.06228727},"labels":[],"label_agreement":null},{"id":"W4286331374","doi":"10.1109/saner53432.2022.00152","title":"Exploring Relevant Artifacts of Release Notes: The Practitioners' Perspective","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Software Analysis, Evolution and Reengineering (SANER)","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software release life cycle; Computer science; Key (lock); Set (abstract data type); Perspective (graphical); Reading (process); Software; Software development; World Wide Web; Software engineering; Software quality; Computer security","score_opus":0.07969162625759513,"score_gpt":0.30130812544751046,"score_spread":0.22161649918991533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286331374","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80921304,0.0065751197,0.13118422,0.017060518,0.00019187461,0.00060489477,0.0003881859,0.00050393597,0.034278236],"genre_scores_gemma":[0.9227993,0.0039798943,0.06586538,0.0010914064,0.00008761768,0.00021115791,0.0004360803,0.0002257203,0.005303482],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.96634936,0.022791222,0.0013625622,0.0020578392,0.0060742632,0.0013647912],"domain_scores_gemma":[0.859935,0.10768355,0.010645749,0.007467689,0.01213888,0.0021291245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026928457,0.0010723921,0.0005296789,0.007601203,0.0037796574,0.011102735,0.002777843,0.003369753,0.0025547226],"category_scores_gemma":[0.07082394,0.0010656158,0.0006960011,0.0059225485,0.0066853077,0.0128086675,0.005098246,0.0027502417,0.00064725697],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011146537,0.0003943623,0.04223878,0.0016854589,0.000044602457,0.0065935734,0.7651602,0.0010884295,0.008522641,0.019473998,0.0035954723,0.15109111],"study_design_scores_gemma":[0.000039852395,0.00035727286,0.033265837,0.0022768318,0.0000967413,0.004087316,0.80625,0.0041351602,0.0063973395,0.018049391,0.12493488,0.0001094165],"about_ca_topic_score_codex":0.0046900045,"about_ca_topic_score_gemma":0.008548738,"teacher_disagreement_score":0.026928457,"about_ca_system_score_codex":0.00347062,"about_ca_system_score_gemma":0.0059069935,"threshold_uncertainty_score":0.14241296},"labels":[],"label_agreement":null},{"id":"W4286331421","doi":"10.1109/saner53432.2022.00090","title":"Toward Understanding the Impact of Refactoring on Program Comprehension","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Software Analysis, Evolution and Reengineering (SANER)","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung","keywords":"Code refactoring; Program comprehension; Computer science; Readability; Software maintenance; Software engineering; Programming language; Maintainability; Commit; Source code; Software; Software system; Database","score_opus":0.09590713737553741,"score_gpt":0.34146979367742475,"score_spread":0.24556265630188734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286331421","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.901341,0.001903945,0.08920192,0.0017200554,0.000016355194,0.0001181198,0.0011404599,0.00054229604,0.0040158373],"genre_scores_gemma":[0.9731674,0.0006364296,0.02465148,0.00012107728,0.000028222854,0.00007779407,0.0008810656,0.00012244984,0.0003140374],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9912334,0.004256844,0.00057705527,0.0014816197,0.0019903595,0.0004607019],"domain_scores_gemma":[0.6950369,0.24867505,0.027711783,0.013365358,0.0136497,0.0015612503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010712212,0.0011587258,0.0008763666,0.005560971,0.00048812485,0.004235115,0.001301222,0.0014901529,0.0016241217],"category_scores_gemma":[0.12601131,0.00080264243,0.0010207633,0.0042468146,0.0018204352,0.01240908,0.0019860202,0.0030230999,0.00043486612],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003844629,0.00070477574,0.7764458,0.00085304014,0.00043641808,0.00018418078,0.006715702,0.036929484,0.01722379,0.0077994564,0.000861745,0.15146117],"study_design_scores_gemma":[0.000021756827,0.00096594845,0.79564667,0.0002813665,0.00024143473,0.00019575187,0.0026894142,0.16250096,0.013836268,0.020692239,0.0027888154,0.00013936571],"about_ca_topic_score_codex":0.008219771,"about_ca_topic_score_gemma":0.0075885844,"teacher_disagreement_score":0.010712212,"about_ca_system_score_codex":0.0013999542,"about_ca_system_score_gemma":0.0017726303,"threshold_uncertainty_score":0.056652248},"labels":[],"label_agreement":null},{"id":"W4286530317","doi":"10.1109/saner53432.2022.00107","title":"On the Importance of Performing App Analysis Within Peer Groups","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Software Analysis, Evolution and Reengineering (SANER)","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kingston Health Sciences Centre; Polytechnique Montréal; Queen's University; Thompson Rivers University","funders":"","keywords":"Perspective (graphical); Computer science; Mobile apps; App store; World Wide Web; Smartphone app; Android app; Internet privacy; Peer-to-peer; Peer review; Focus group; Android (operating system); Context (archaeology); Data science; Artificial intelligence; Business; Marketing","score_opus":0.025822111897008344,"score_gpt":0.27341101624327296,"score_spread":0.24758890434626463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286530317","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69894004,0.0046138023,0.2568505,0.0064536356,0.0008280391,0.0015516209,0.0009904075,0.0022824015,0.027489474],"genre_scores_gemma":[0.9108264,0.00051199907,0.08499335,0.00040289306,0.00034284955,0.0004168889,0.00046707224,0.00035745907,0.0016810698],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.94581723,0.03316332,0.002583046,0.006090114,0.010938718,0.0014076668],"domain_scores_gemma":[0.7272053,0.1814981,0.019456014,0.024275864,0.042516746,0.005047987],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.037349913,0.00127289,0.0020894024,0.0046798475,0.0031368986,0.0069124973,0.0015702125,0.0019368172,0.0023267195],"category_scores_gemma":[0.1794826,0.00071354775,0.00089615554,0.0030600384,0.0021095765,0.01070519,0.004225233,0.0021500764,0.0020155092],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001406141,0.0010194156,0.36174718,0.0017048668,0.0011615785,0.0010491556,0.023433141,0.0082621975,0.029298445,0.0105160335,0.014554528,0.5458473],"study_design_scores_gemma":[0.00027091254,0.0034467557,0.5321232,0.0009806082,0.001286394,0.0035897386,0.04757243,0.2201065,0.033627763,0.08075001,0.07558503,0.0006606908],"about_ca_topic_score_codex":0.004884277,"about_ca_topic_score_gemma":0.007161713,"teacher_disagreement_score":0.037349913,"about_ca_system_score_codex":0.0010322708,"about_ca_system_score_gemma":0.0027899116,"threshold_uncertainty_score":0.19752759},"labels":[],"label_agreement":null},{"id":"W4286531983","doi":"10.1109/saner53432.2022.00094","title":"Automatic Detection and Analysis of Technical Debts in Peer-Review Documentation of R Packages","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Software Analysis, Evolution and Reengineering (SANER)","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Technical debt; Documentation; Computer science; Suite; Usability; Empirical research; Popularity; Technical documentation; Artificial intelligence; Software engineering; World Wide Web; Software; Machine learning; Software development; Programming language; Operating system; Statistics","score_opus":0.025277533824908677,"score_gpt":0.31516048281083275,"score_spread":0.2898829489859241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286531983","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7436601,0.016207138,0.14206465,0.008515767,0.0019964816,0.0020581447,0.043778464,0.026228704,0.015490486],"genre_scores_gemma":[0.7143373,0.003649603,0.21879385,0.0017374094,0.0013600081,0.0022933849,0.04044946,0.004923287,0.012455624],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9450749,0.018692505,0.009584921,0.008282197,0.017342618,0.0010228697],"domain_scores_gemma":[0.5177994,0.25764927,0.11518468,0.027417246,0.07889153,0.0030579413],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.05124974,0.001300794,0.0013700129,0.018738704,0.0018539946,0.0036064482,0.0023504011,0.0015287466,0.0033382743],"category_scores_gemma":[0.30696413,0.00082043966,0.0010054953,0.010201569,0.001037183,0.0043089967,0.0033119156,0.0019745748,0.004727364],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082794414,0.0002641737,0.32028785,0.009832211,0.0006896919,0.0021369103,0.011470156,0.004624975,0.021392742,0.0055452907,0.17407396,0.4488541],"study_design_scores_gemma":[0.00023557738,0.0005330356,0.430801,0.0041813436,0.0007854792,0.0053864755,0.008991593,0.1491401,0.03666835,0.011246526,0.35116732,0.0008631402],"about_ca_topic_score_codex":0.0036767581,"about_ca_topic_score_gemma":0.008546164,"teacher_disagreement_score":0.94875026,"about_ca_system_score_codex":0.001715026,"about_ca_system_score_gemma":0.0050808224,"threshold_uncertainty_score":0.27103776},"labels":[],"label_agreement":null},{"id":"W4286532194","doi":"10.1109/saner53432.2022.00014","title":"Do Developers Refactor Data Access Code? An Empirical Study","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Software Analysis, Evolution and Reengineering (SANER)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Fonds de Recherche du Québec - Santé; Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Maintainability; Data access; Program comprehension; Software; Software engineering; Database; Software system; Programming language","score_opus":0.14237056991475835,"score_gpt":0.3996569684777112,"score_spread":0.2572863985629529,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286532194","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9971384,0.0003623195,0.0006375984,0.0004739754,0.000006268773,0.00008741748,0.00017688179,0.000017240516,0.0010998993],"genre_scores_gemma":[0.99741334,0.0003871184,0.001059177,0.00021325704,0.000013718596,0.00012774322,0.0002355157,0.000017870732,0.00053226046],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9773048,0.009028743,0.0026932377,0.0027624634,0.006803165,0.0014076026],"domain_scores_gemma":[0.48122707,0.3799474,0.07916533,0.013856767,0.040880807,0.0049225036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028266937,0.00040631735,0.0004531764,0.0042136847,0.0013743194,0.0021992382,0.0014267816,0.0016011175,0.002080041],"category_scores_gemma":[0.21677278,0.0007313782,0.00035431457,0.003686266,0.001999026,0.0049608136,0.0018958008,0.0023284801,0.0006799164],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020793648,0.0012545796,0.89145845,0.00042379068,0.000065049164,0.0006579363,0.065155275,0.00014871721,0.00062973297,0.00045491414,0.0015562223,0.037987325],"study_design_scores_gemma":[0.00008569745,0.00087460934,0.9066766,0.0005305586,0.000079635334,0.0014087782,0.076586716,0.0021213368,0.0010597579,0.0005613652,0.009945498,0.00006944427],"about_ca_topic_score_codex":0.0058317026,"about_ca_topic_score_gemma":0.008701372,"teacher_disagreement_score":0.028266937,"about_ca_system_score_codex":0.0020054337,"about_ca_system_score_gemma":0.0026444774,"threshold_uncertainty_score":0.14949161},"labels":[],"label_agreement":null},{"id":"W4287017899","doi":"10.5281/zenodo.5663903","title":"Mapping breakpoint types: an exploratory study","year":2021,"lang":"en","type":"dataset","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Breakpoint; Debugging; Computer science; Documentation; Programming language; Biology","score_opus":0.08246056451405347,"score_gpt":0.21053348792486504,"score_spread":0.12807292341081156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287017899","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00928548,0.00057681475,0.0009718634,0.0002709325,0.000050769402,0.00014603087,0.9847794,0.0005993121,0.0033193745],"genre_scores_gemma":[0.006600563,0.00027424542,0.0026467445,0.00014868335,0.00001675827,0.00057635293,0.9879732,0.00014292917,0.0016205367],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966073,0.000771639,0.00074245915,0.0007155221,0.00084036426,0.00032280403],"domain_scores_gemma":[0.98453945,0.0073359464,0.0015542468,0.0020362702,0.003901353,0.00063264574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029124343,0.0010036703,0.0007735272,0.012243408,0.0012728815,0.0023266464,0.001806939,0.0015824686,0.015381711],"category_scores_gemma":[0.018178511,0.00039504402,0.00073352427,0.014271163,0.00048473882,0.00184362,0.0027146356,0.0014056774,0.018308401],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042445443,0.00026244414,0.030276844,0.005504702,0.00009520674,0.0005645021,0.0011934319,0.00057540456,0.001459676,0.0037153675,0.9232976,0.032630406],"study_design_scores_gemma":[0.0002312922,0.000066057975,0.03645621,0.00115438,0.00006057163,0.0005024044,0.0017436729,0.00106596,0.001901766,0.0019221879,0.9548361,0.000059433503],"about_ca_topic_score_codex":0.0111884605,"about_ca_topic_score_gemma":0.026452042,"teacher_disagreement_score":0.015381711,"about_ca_system_score_codex":0.0017070759,"about_ca_system_score_gemma":0.0021966908,"threshold_uncertainty_score":0.051456988},"labels":[],"label_agreement":null},{"id":"W4287064711","doi":"10.48550/arxiv.2107.13708","title":"Learning how to listen: Automatically finding bug patterns in event-driven JavaScript APIs","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"JavaScript; Computer science; Event (particle physics); Set (abstract data type); Code (set theory); Ajax; Java; Source code; Programming language; Artificial intelligence; Data mining; Natural language processing; Web service","score_opus":0.06200842605263655,"score_gpt":0.21093014267509536,"score_spread":0.1489217166224588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287064711","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6065153,0.0014643233,0.33502647,0.00073828234,0.00013557087,0.00038730868,0.0057086693,0.04844021,0.0015839151],"genre_scores_gemma":[0.6871973,0.00038980445,0.29707214,0.00026928037,0.0000625028,0.00027595734,0.011545935,0.0012812744,0.001905748],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99737763,0.00042846033,0.0002721982,0.0011681238,0.0005838703,0.00016971884],"domain_scores_gemma":[0.9886952,0.0066771866,0.0018168946,0.0011460605,0.0013023925,0.0003623397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018090457,0.0018526404,0.00085011724,0.0040581464,0.00044082344,0.0012330613,0.0018668959,0.0013830839,0.00052145135],"category_scores_gemma":[0.014024967,0.0006331804,0.0012072903,0.001928888,0.0006507793,0.0029699323,0.0013915933,0.0013239868,0.0008903769],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005699023,0.0009722285,0.22583,0.0010418165,0.0003263414,0.001661988,0.0023424493,0.02277417,0.034250546,0.0014352697,0.02027034,0.68852496],"study_design_scores_gemma":[0.00015003205,0.00039079884,0.06068198,0.00013955428,0.00024795742,0.001408284,0.0010879475,0.8808793,0.036060743,0.0075664753,0.01125951,0.0001274958],"about_ca_topic_score_codex":0.0040198923,"about_ca_topic_score_gemma":0.0067639663,"teacher_disagreement_score":0.0040581464,"about_ca_system_score_codex":0.00048219084,"about_ca_system_score_gemma":0.0010908502,"threshold_uncertainty_score":0.009567261},"labels":[],"label_agreement":null},{"id":"W4287279937","doi":"10.5281/zenodo.4606679","title":"How Effective is Continuous Integration in Indicating Single-Statement Bugs?","year":2021,"lang":"en","type":"paratext","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Concordia University","funders":"","keywords":"Computer science; Statement (logic); Software bug; Programming language; Software","score_opus":0.02616547045364239,"score_gpt":0.2530509519036387,"score_spread":0.2268854814499963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287279937","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92956764,0.0027365563,0.054835867,0.0007053254,0.000088761044,0.00015507033,0.0005329923,0.0062898044,0.005087999],"genre_scores_gemma":[0.9793168,0.00033187753,0.018903792,0.00013451587,0.000030405125,0.000038033348,0.0003720785,0.00033049745,0.00054205814],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98511547,0.004007801,0.0011261303,0.0021831812,0.0064454973,0.0011218747],"domain_scores_gemma":[0.777279,0.16057618,0.031111682,0.0128492685,0.015042503,0.0031412798],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01111417,0.000972354,0.00093204214,0.005699651,0.0006227277,0.0026758257,0.0016119987,0.0016504138,0.0017798634],"category_scores_gemma":[0.10598829,0.00071365386,0.00051558623,0.003219274,0.0019545215,0.004739497,0.0020361478,0.0013678214,0.00090219895],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011640243,0.00040090617,0.5648513,0.0009466963,0.00026114244,0.00068779563,0.0027553812,0.0104577225,0.040262517,0.0010258206,0.0032322651,0.37395453],"study_design_scores_gemma":[0.00011690112,0.0042055757,0.7675577,0.00071468775,0.00066525437,0.0033786916,0.0032724084,0.14025791,0.06672734,0.003779941,0.008929428,0.00039413475],"about_ca_topic_score_codex":0.0049232603,"about_ca_topic_score_gemma":0.005316624,"teacher_disagreement_score":0.01111417,"about_ca_system_score_codex":0.0007001921,"about_ca_system_score_gemma":0.0011543438,"threshold_uncertainty_score":0.058778048},"labels":[],"label_agreement":null},{"id":"W4287383862","doi":"10.1007/978-3-031-10548-7_31","title":"Security Evaluation Criteria of Open-Source Libraries","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University of Edmonton","funders":"","keywords":"Computer science; Source code; Open source; Software engineering; Software; Code review; Process (computing); Code (set theory); Open source software; Static program analysis; Checklist; World Wide Web; Computer security; Database; Software development; Operating system; Programming language","score_opus":0.0364567715616366,"score_gpt":0.30967840584871315,"score_spread":0.2732216342870766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287383862","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6532708,0.008149189,0.2602791,0.0011612355,0.00022227844,0.00080932426,0.0010858094,0.0013508602,0.073671475],"genre_scores_gemma":[0.9476263,0.0005928396,0.04553063,0.00004686591,0.00009199013,0.00012779923,0.0008253692,0.00009475809,0.005063479],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99067163,0.002472624,0.0007227714,0.0003664616,0.0052525704,0.0005139489],"domain_scores_gemma":[0.959746,0.022846226,0.002504203,0.001538087,0.011701562,0.0016639087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007948666,0.0008281522,0.00080827455,0.009483308,0.0010878288,0.004234721,0.0009079168,0.0008439626,0.0039108223],"category_scores_gemma":[0.030770915,0.00022412362,0.0007689053,0.002912767,0.0010093907,0.0032853496,0.0015922906,0.0006744691,0.00061689224],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003991862,0.000826631,0.057408426,0.0015404491,0.00026208133,0.0005625614,0.0009770422,0.060070638,0.03337691,0.12415221,0.015298109,0.701533],"study_design_scores_gemma":[0.00022436866,0.0023671775,0.043486368,0.000869145,0.00051682536,0.0012740915,0.0017107587,0.79337007,0.05325136,0.08282497,0.019925907,0.0001789526],"about_ca_topic_score_codex":0.0015698405,"about_ca_topic_score_gemma":0.0014660235,"teacher_disagreement_score":0.009483308,"about_ca_system_score_codex":0.0023428437,"about_ca_system_score_gemma":0.0015832983,"threshold_uncertainty_score":0.04203707},"labels":[],"label_agreement":null},{"id":"W4287558158","doi":"","title":"Analysing Microsoft Access Projects:Building a model in a Partially Observable Domain","year":2020,"lang":"en","type":"article","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Berger (Canada)","funders":"","keywords":"Computer science; Observable; Domain (mathematical analysis); Microsoft Office; Domain model; Software engineering; Programming language; Domain knowledge; Mathematics","score_opus":0.040381571729665035,"score_gpt":0.2729170395765636,"score_spread":0.23253546784689857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287558158","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32769465,0.00016609483,0.66575676,0.00059978187,0.000014824625,0.00027800907,0.0011608085,0.0017703472,0.0025587238],"genre_scores_gemma":[0.61643904,0.00021893521,0.37911934,0.0000528431,0.000013057636,0.0003874076,0.0024763877,0.0003289567,0.0009639804],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99458057,0.0024192033,0.00051357475,0.0010287528,0.0011791886,0.00027863175],"domain_scores_gemma":[0.9773751,0.013859567,0.0021786669,0.004707011,0.001520993,0.00035870212],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004859499,0.00080837624,0.0005705305,0.0021238683,0.00073618116,0.0031928683,0.00141604,0.0014880097,0.0012026854],"category_scores_gemma":[0.023077501,0.00074084057,0.0014894514,0.0018035355,0.0017992964,0.0062457947,0.002812411,0.0015027442,0.0003616271],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084390544,0.0009710231,0.12884991,0.0013489706,0.0002825194,0.0026832563,0.01739801,0.52272964,0.026404133,0.13997751,0.0025807882,0.15593033],"study_design_scores_gemma":[0.000062348554,0.0003208604,0.01043364,0.00024241027,0.0000971413,0.00039810347,0.002891966,0.8917216,0.012098947,0.06912687,0.012528143,0.00007801936],"about_ca_topic_score_codex":0.010272009,"about_ca_topic_score_gemma":0.007028149,"teacher_disagreement_score":0.010272009,"about_ca_system_score_codex":0.0011908142,"about_ca_system_score_gemma":0.0028568343,"threshold_uncertainty_score":0.025699735},"labels":[],"label_agreement":null},{"id":"W4287668178","doi":"10.48550/arxiv.2009.09930","title":"AOBTM: Adaptive Online Biterm Topic Modeling for Version Sensitive\\n Short-texts Analysis","year":2020,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; Kelowna General Hospital","funders":"","keywords":"Computer science; Topic model; Inference; Online algorithm; Mobile apps; Information retrieval; Word (group theory); Data mining; Machine learning; Data science; Artificial intelligence; World Wide Web; Algorithm","score_opus":0.14680901011834516,"score_gpt":0.23638759999836267,"score_spread":0.08957858988001752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287668178","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01781066,0.001900225,0.97170466,0.000540428,0.00025250166,0.00035699434,0.0020362556,0.0045292783,0.0008690459],"genre_scores_gemma":[0.35993022,0.0028301212,0.6049336,0.00096752256,0.0015697307,0.002443685,0.016594201,0.0010469698,0.009683866],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970277,0.0010630523,0.00029751298,0.0009394856,0.00047441706,0.00019777684],"domain_scores_gemma":[0.9940574,0.00396314,0.00046613705,0.0006113906,0.00069746864,0.00020432202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037737926,0.001979583,0.0020191283,0.003408377,0.00090401684,0.0018985367,0.0034706923,0.0022497175,0.0034816754],"category_scores_gemma":[0.012658607,0.00107586,0.0026817406,0.0032755255,0.0007227968,0.0036077232,0.0026962417,0.0037404117,0.0032662272],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011710511,0.0006750919,0.009307012,0.0009825688,0.0006610197,0.00047315776,0.0014788642,0.14538491,0.011435953,0.01680231,0.024538113,0.78708994],"study_design_scores_gemma":[0.00003609079,0.000058472295,0.0009965914,0.000024543291,0.00004883787,0.00006710253,0.00007219146,0.9848461,0.0010622893,0.008889064,0.003873273,0.000025399195],"about_ca_topic_score_codex":0.009888488,"about_ca_topic_score_gemma":0.012393679,"teacher_disagreement_score":0.009888488,"about_ca_system_score_codex":0.0009986279,"about_ca_system_score_gemma":0.0016714123,"threshold_uncertainty_score":0.01995796},"labels":[],"label_agreement":null},{"id":"W4287776252","doi":"10.5281/zenodo.3839075","title":"Characterizing Task-Relevant Information in Natural Language Software Artifacts","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Task (project management); Natural language; Natural (archaeology); Software; Natural language processing; Human–computer interaction; Artificial intelligence; Programming language; Engineering; Geography; Systems engineering","score_opus":0.018120083661019887,"score_gpt":0.22331582101500974,"score_spread":0.20519573735398985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287776252","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5025124,0.0018869317,0.32079634,0.0031733462,0.0009137184,0.002677332,0.12293008,0.010361337,0.034748487],"genre_scores_gemma":[0.5224282,0.0006311604,0.31558585,0.0009654378,0.00025817746,0.002611482,0.1471693,0.0013281051,0.009022213],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99732816,0.0011335289,0.00028362905,0.00056200934,0.00058619346,0.00010631324],"domain_scores_gemma":[0.92200005,0.062978104,0.0037783412,0.003493468,0.007019519,0.00073060667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003176958,0.0006242347,0.00029223555,0.0023709375,0.00073287275,0.0019451253,0.00072742696,0.00094262854,0.027658112],"category_scores_gemma":[0.04682869,0.00022820968,0.00039624583,0.0014967688,0.00039514984,0.0020197923,0.0013275936,0.000661513,0.006423256],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020940532,0.0023479408,0.057172272,0.009763328,0.0001970292,0.0020150398,0.009998286,0.006433312,0.11128376,0.02069013,0.27218953,0.5058154],"study_design_scores_gemma":[0.00057876867,0.0019335083,0.23847066,0.0025778662,0.0003608122,0.0044632256,0.010785172,0.13230859,0.122259706,0.05639923,0.42928934,0.00057319493],"about_ca_topic_score_codex":0.0018428338,"about_ca_topic_score_gemma":0.0041285143,"teacher_disagreement_score":0.027658112,"about_ca_system_score_codex":0.0007457485,"about_ca_system_score_gemma":0.0010123001,"threshold_uncertainty_score":0.0925256},"labels":[],"label_agreement":null},{"id":"W4288281409","doi":"10.48550/arxiv.1907.07803","title":"Syntax and Stack Overflow: A methodology for extracting a corpus of\\n syntax errors and fixes","year":2019,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Syntax error; Computer science; Python (programming language); Syntax; Abstract syntax tree; Parsing; Programming language; Abstract syntax; Source code; Natural language processing; Artificial intelligence","score_opus":0.17718333111090953,"score_gpt":0.2577906096931637,"score_spread":0.08060727858225417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4288281409","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.097652115,0.0017231458,0.16663067,0.0016480978,0.0007173632,0.003716281,0.65848076,0.049733236,0.019698266],"genre_scores_gemma":[0.051352575,0.00054445234,0.29503402,0.0006266602,0.00015354948,0.006139663,0.6329473,0.0059215766,0.0072801565],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.985182,0.0024004092,0.0037005357,0.003959078,0.0040984717,0.0006594274],"domain_scores_gemma":[0.96411884,0.013561928,0.005459573,0.008209655,0.007831549,0.0008183516],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008007989,0.0022847403,0.0010520116,0.021308811,0.0023646455,0.0033566516,0.0028572308,0.0026684287,0.008875805],"category_scores_gemma":[0.04307223,0.0012599623,0.0020801162,0.014406371,0.0025067565,0.006334919,0.007834288,0.0033545357,0.0123014],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057425327,0.0005089147,0.059419967,0.008846986,0.00039252854,0.0022035483,0.015696382,0.0041304496,0.027338123,0.020337922,0.5186789,0.34187207],"study_design_scores_gemma":[0.00016342475,0.00019512208,0.09967879,0.0010705413,0.00020186423,0.001691092,0.0050604083,0.019174682,0.028455606,0.01846845,0.82540125,0.00043876597],"about_ca_topic_score_codex":0.014920122,"about_ca_topic_score_gemma":0.02821362,"teacher_disagreement_score":0.021308811,"about_ca_system_score_codex":0.002373378,"about_ca_system_score_gemma":0.007340484,"threshold_uncertainty_score":0.04235077},"labels":[],"label_agreement":null},{"id":"W4288321959","doi":"10.48550/arxiv.1906.07812","title":"Debunking the Myth that Upfront Requirements are Infeasible for\\n Scientific Computing Software","year":2019,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Software engineering; Documentation; Software requirements; Software requirements specification; Requirements analysis; Traceability; Software; Software development; Requirements traceability; Software design; Systems engineering; Requirement; Programming language; Engineering","score_opus":0.1777583518229039,"score_gpt":0.23805681266302814,"score_spread":0.060298460840124246,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4288321959","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030767828,0.0052349344,0.497081,0.37161884,0.0042654653,0.00009668991,0.00016316633,0.0024093653,0.08836268],"genre_scores_gemma":[0.45085126,0.010289685,0.41622654,0.065188564,0.0042774873,0.00041900075,0.00032204494,0.0033759773,0.04904947],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97667116,0.010462294,0.0011575829,0.0022487224,0.0087860115,0.0006743118],"domain_scores_gemma":[0.90100867,0.059951514,0.0042357803,0.022708388,0.010615508,0.0014801488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025042681,0.0007394954,0.00045331233,0.001450166,0.0042582974,0.009038114,0.0026365232,0.0047687255,0.007034454],"category_scores_gemma":[0.083464265,0.00084816,0.00080795656,0.00086997857,0.022619057,0.02548191,0.00677318,0.015029662,0.0031455962],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007042418,0.000037606012,0.0006056574,0.00017196588,0.000021073265,0.00016966862,0.006398509,0.0015360024,0.0013829842,0.8657303,0.042160563,0.08171536],"study_design_scores_gemma":[0.000027084137,0.000053170712,0.0003211783,0.00035396186,0.000015651807,0.00040848056,0.0017796783,0.0035026725,0.0021843903,0.7053361,0.28594866,0.00006905098],"about_ca_topic_score_codex":0.002232639,"about_ca_topic_score_gemma":0.0020536594,"teacher_disagreement_score":0.025042681,"about_ca_system_score_codex":0.0036620868,"about_ca_system_score_gemma":0.003713597,"threshold_uncertainty_score":0.13243997},"labels":[],"label_agreement":null},{"id":"W4288562799","doi":"","title":"GUI Migration using MDE from GWT to Angular 6:An Industrial Case","year":2019,"lang":"en","type":"article","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Berger (Canada)","funders":"","keywords":"Computer science","score_opus":0.029051185038978092,"score_gpt":0.25549277671618825,"score_spread":0.22644159167721017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4288562799","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7160799,0.00036328819,0.19349724,0.0012717226,0.00029885495,0.000370687,0.0006836984,0.05351333,0.033921257],"genre_scores_gemma":[0.8934624,0.00012831448,0.09098635,0.0001899665,0.000018848606,0.00007704587,0.000660245,0.0033480192,0.011128765],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990969,0.00025688414,0.00005559458,0.00013718373,0.00029049662,0.00016302754],"domain_scores_gemma":[0.9969093,0.0012928735,0.00010568039,0.0010751402,0.00035143914,0.00026563564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012851397,0.00071523554,0.0004738263,0.00073981856,0.0005721506,0.0014215173,0.0020051827,0.0014192902,0.010637857],"category_scores_gemma":[0.0063441675,0.00043613833,0.0006850332,0.0007699626,0.0008234462,0.0012633824,0.0015676835,0.001484772,0.0019598186],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0064498917,0.0026194109,0.02524198,0.0011326742,0.00024048562,0.014804887,0.0049422765,0.18680306,0.097437166,0.028958961,0.037749022,0.5936202],"study_design_scores_gemma":[0.00096340664,0.0016010737,0.016896276,0.00020350696,0.00022901315,0.002950505,0.0016962952,0.7762493,0.11247171,0.009334879,0.07720508,0.00019902168],"about_ca_topic_score_codex":0.0056844787,"about_ca_topic_score_gemma":0.0043453746,"teacher_disagreement_score":0.010637857,"about_ca_system_score_codex":0.00049812475,"about_ca_system_score_gemma":0.00067957694,"threshold_uncertainty_score":0.035587132},"labels":[],"label_agreement":null},{"id":"W4289170436","doi":"10.1142/s0218194022500498","title":"A Semantic Web-Enabled Approach for Dependency Management","year":2022,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software engineering; Artifact (error); Knowledge base; Software development; Software; World Wide Web; Data science; Programming language; Artificial intelligence","score_opus":0.01051651063017205,"score_gpt":0.24127503706056194,"score_spread":0.23075852643038988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289170436","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027829057,0.00019222543,0.98766196,0.0011464946,0.000044305,0.00010018155,0.00033691456,0.00083704304,0.006897845],"genre_scores_gemma":[0.11229721,0.00085706706,0.8798357,0.0003597687,0.00008022829,0.00033008924,0.0015438813,0.00029517984,0.0044009257],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99612385,0.0012919314,0.00044678515,0.0007102652,0.001172231,0.0002549943],"domain_scores_gemma":[0.996633,0.0010304442,0.0003373643,0.001186559,0.00062773126,0.00018477473],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004042274,0.0011397825,0.0007743934,0.007459305,0.0022781477,0.006372332,0.0028772028,0.002232984,0.0029841496],"category_scores_gemma":[0.005499658,0.0012132811,0.0036839934,0.006191065,0.0033625704,0.0135945175,0.0055548805,0.0034074038,0.001116972],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003814,0.000112612724,0.0015994593,0.00021654238,0.0001306571,0.0007764034,0.0016635014,0.036335506,0.0017292341,0.8959796,0.0047739926,0.056644432],"study_design_scores_gemma":[0.000019570054,0.000026193544,0.00050530495,0.00021978129,0.000115361996,0.00045986322,0.00072461116,0.23625642,0.0027274082,0.6622991,0.09659149,0.000054898603],"about_ca_topic_score_codex":0.014082365,"about_ca_topic_score_gemma":0.016873058,"teacher_disagreement_score":0.014082365,"about_ca_system_score_codex":0.0027611207,"about_ca_system_score_gemma":0.0041424944,"threshold_uncertainty_score":0.028000832},"labels":[],"label_agreement":null},{"id":"W4289518662","doi":"10.1002/smr.2499","title":"Release conventions of open‐source software: An exploratory study","year":2022,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Canada First Research Excellence Fund","keywords":"Computer science; Software development; Software engineering; Software; Social software engineering; Personal software process; Software peer review; Software release life cycle; Interview; Coding (social sciences); Software development process; Data science; Software construction","score_opus":0.029983820538575456,"score_gpt":0.309171393674136,"score_spread":0.27918757313556053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289518662","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99818146,0.0000259374,0.0010525602,0.000040799925,0.0000022176919,0.000054280838,0.00016725335,0.000018863651,0.00045660642],"genre_scores_gemma":[0.99546045,0.00005464544,0.003250484,0.000029304154,0.0000055723353,0.000113154034,0.0006308551,0.00003510266,0.0004203577],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9927859,0.00461598,0.00038457307,0.0007773815,0.0011401067,0.00029614862],"domain_scores_gemma":[0.9201726,0.063617,0.006165068,0.004859933,0.003634349,0.0015510564],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0067746327,0.00026969044,0.0002934576,0.002072912,0.0012770322,0.0017582404,0.0010286439,0.00071183825,0.00069543225],"category_scores_gemma":[0.039054736,0.00029958322,0.00032687272,0.0020173169,0.0014366364,0.0026127035,0.0017934231,0.001152578,0.0002720778],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007869401,0.0028659822,0.47321764,0.00085853983,0.00010973589,0.00337697,0.38928112,0.0024347473,0.017577777,0.0026438471,0.0031053107,0.10374133],"study_design_scores_gemma":[0.00007916058,0.0016665227,0.7098705,0.0004007608,0.00007233731,0.0018706527,0.22773525,0.016437616,0.010675113,0.0026040785,0.028348014,0.00023999106],"about_ca_topic_score_codex":0.002263398,"about_ca_topic_score_gemma":0.0055328608,"teacher_disagreement_score":0.9932254,"about_ca_system_score_codex":0.00087792735,"about_ca_system_score_gemma":0.0005473097,"threshold_uncertainty_score":0.035828114},"labels":[],"label_agreement":null},{"id":"W4289730357","doi":"10.48550/arxiv.1808.00594","title":"Improving IR-Based Bug Localization with Context-Aware Query\\n Reformulation","year":2018,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Information retrieval; Context (archaeology); Query expansion; Baseline (sea); State (computer science); Query language; Data mining; Natural language processing; Programming language","score_opus":0.05079169819544985,"score_gpt":0.19202318096687193,"score_spread":0.14123148277142208,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289730357","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12638548,0.007451361,0.7737613,0.0014848615,0.00032018498,0.0005910219,0.0018457656,0.08329123,0.004868752],"genre_scores_gemma":[0.3506828,0.0016240795,0.6332888,0.0009126648,0.0003102514,0.00023651082,0.005676731,0.0016848294,0.005583285],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9967673,0.0007276654,0.0003210284,0.0008110142,0.0011501289,0.00022283543],"domain_scores_gemma":[0.9932407,0.0026095596,0.0008709401,0.0013973202,0.0017253375,0.00015601238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020888855,0.0019757638,0.0021301107,0.0058893263,0.00083336496,0.0014949505,0.0022887157,0.0013802915,0.0035732042],"category_scores_gemma":[0.010724256,0.000501374,0.0015841366,0.0033059006,0.0008603582,0.0037264766,0.0021336845,0.0013933013,0.0027033314],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039086037,0.00047338285,0.006272988,0.001064622,0.0001357182,0.0004977084,0.0011327331,0.012505243,0.10414078,0.0035796883,0.02746134,0.8423449],"study_design_scores_gemma":[0.00037257918,0.0015564357,0.014292693,0.00020519827,0.0009841057,0.003388894,0.001762021,0.70686626,0.20439766,0.01097449,0.054838184,0.00036143372],"about_ca_topic_score_codex":0.010188154,"about_ca_topic_score_gemma":0.008350782,"teacher_disagreement_score":0.010188154,"about_ca_system_score_codex":0.00096210255,"about_ca_system_score_gemma":0.0019672061,"threshold_uncertainty_score":0.020257711},"labels":[],"label_agreement":null},{"id":"W4289754143","doi":"10.48550/arxiv.1807.04475","title":"STRICT: Information Retrieval Based Search Term Identification for\\n Concept Location","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Task (project management); Information retrieval; Term (time); Identification (biology); Software; Baseline (sea); Natural language; Domain (mathematical analysis); Software maintenance; Source code; Quality (philosophy); Software development; Natural language processing; Artificial intelligence; Data mining; Programming language","score_opus":0.06533597192685373,"score_gpt":0.22026414261055768,"score_spread":0.15492817068370396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289754143","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1716897,0.009819311,0.7108492,0.0019806596,0.0012106167,0.0018740555,0.016721696,0.07088579,0.014968846],"genre_scores_gemma":[0.2374735,0.0018062682,0.712128,0.00055877835,0.00053823146,0.0007407327,0.029379109,0.0013573258,0.016018009],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99778146,0.00050020794,0.00025288042,0.000404943,0.0008851397,0.00017532092],"domain_scores_gemma":[0.9955987,0.0020438011,0.0005047279,0.00059299363,0.0010237651,0.00023592137],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00148888,0.0017197446,0.0015133611,0.012106884,0.0010696092,0.0015891902,0.001747924,0.0014927016,0.006778294],"category_scores_gemma":[0.0074845576,0.00044652662,0.0011521758,0.0057386067,0.00065558765,0.004257388,0.0018776917,0.0011923527,0.0065283016],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009079816,0.0007090925,0.009561795,0.0020448328,0.00024617533,0.00046720504,0.0006636097,0.0041829636,0.08267032,0.0061605847,0.09444965,0.7979358],"study_design_scores_gemma":[0.0008186127,0.002181553,0.028191462,0.00037982033,0.0007340027,0.004160255,0.0018779952,0.59976363,0.16404809,0.030822985,0.16646896,0.0005526253],"about_ca_topic_score_codex":0.0069606504,"about_ca_topic_score_gemma":0.016083991,"teacher_disagreement_score":0.012106884,"about_ca_system_score_codex":0.00097642106,"about_ca_system_score_gemma":0.0025462129,"threshold_uncertainty_score":0.022675693},"labels":[],"label_agreement":null},{"id":"W4289926492","doi":"10.1109/icaica54878.2022.9844464","title":"i-DARTS: Improving differentiable architecture search by using graph and few-shot learning","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Artificial Intelligence and Computer Applications (ICAICA)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Differentiable function; Computer science; Architecture; Graph; Shot (pellet); Artificial intelligence; Theoretical computer science; Mathematics","score_opus":0.07899306736135765,"score_gpt":0.3208955494775624,"score_spread":0.24190248211620474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289926492","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10276427,0.0021327892,0.86466783,0.0007995515,0.00039212956,0.0002935554,0.0005180879,0.021088528,0.0073433127],"genre_scores_gemma":[0.5105386,0.00051195116,0.47595993,0.00084650697,0.00015705923,0.00024186114,0.0026875164,0.0010353916,0.008021216],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99942374,0.00012343729,0.000035238718,0.00019374347,0.00015662303,0.00006721366],"domain_scores_gemma":[0.99901223,0.00033746668,0.00007649135,0.00026797847,0.0002061341,0.000099674515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000904776,0.0021573019,0.0017334265,0.0016345583,0.00068815146,0.0009450287,0.0030411035,0.0019698932,0.003818816],"category_scores_gemma":[0.0037421177,0.00054224156,0.0011488594,0.0011703199,0.00077312184,0.0027936485,0.001317119,0.002199578,0.0016750251],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024050161,0.00035601438,0.002488449,0.00022838045,0.00031948584,0.00020901332,0.00014667987,0.48340014,0.008156886,0.010015717,0.022261875,0.47217682],"study_design_scores_gemma":[0.00002821782,0.00007883316,0.0001456104,0.0000063705684,0.000016879689,0.000030351037,0.000016750648,0.99311405,0.0010357251,0.0046587344,0.00086095906,0.000007500718],"about_ca_topic_score_codex":0.011774253,"about_ca_topic_score_gemma":0.021849474,"teacher_disagreement_score":0.011774253,"about_ca_system_score_codex":0.0014353914,"about_ca_system_score_gemma":0.0017550297,"threshold_uncertainty_score":0.023411393},"labels":[],"label_agreement":null},{"id":"W4291414058","doi":"10.3390/software1030014","title":"An Automated Tool for Upgrading Fortran Codes","year":2022,"lang":"en","type":"article","venue":"Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; British Columbia Institute of Technology; Langara College","funders":"","keywords":"Python (programming language); Computer science; Fortran; Code refactoring; Software portability; Programming language; Software; Coding (social sciences); Software engineering","score_opus":0.018628075643235414,"score_gpt":0.3013823638085817,"score_spread":0.28275428816534626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4291414058","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013705149,0.00019139358,0.55436355,0.00013626595,0.00013126568,0.00044204286,0.0020159776,0.42403895,0.00497533],"genre_scores_gemma":[0.0833702,0.000229285,0.8756835,0.00022152074,0.00006206433,0.0005118092,0.006236236,0.023241762,0.010443581],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980792,0.0002583882,0.00021305737,0.00044093662,0.0008808625,0.00012764211],"domain_scores_gemma":[0.9940849,0.002630943,0.0005272165,0.0012512709,0.001323323,0.00018234197],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016724803,0.0013846008,0.0006775787,0.0030586438,0.0008170514,0.0012746941,0.0017792425,0.00067882315,0.01771238],"category_scores_gemma":[0.010402378,0.0009827563,0.0008123554,0.0011993696,0.00058602594,0.0016375656,0.001762056,0.0013547803,0.008376298],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032759592,0.00022831015,0.0047365404,0.0006035346,0.000056293255,0.0007159605,0.00054255914,0.00532161,0.0396288,0.0042493194,0.070178494,0.873411],"study_design_scores_gemma":[0.0004444715,0.0004434809,0.013303731,0.000509687,0.00013220646,0.0039323093,0.00028807836,0.2675993,0.302093,0.011577638,0.39924178,0.00043427778],"about_ca_topic_score_codex":0.0017434186,"about_ca_topic_score_gemma":0.0018356433,"teacher_disagreement_score":0.01771238,"about_ca_system_score_codex":0.00060607365,"about_ca_system_score_gemma":0.0015836239,"threshold_uncertainty_score":0.05925381},"labels":[],"label_agreement":null},{"id":"W4292291355","doi":"10.1007/s10664-022-10193-8","title":"Revisiting the debate: Are code metrics useful for measuring maintenance effort?","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Science Foundation of Sri Lanka; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software maintenance; Code (set theory); Context (archaeology); Source code; Java; Software metric; Code review; Granularity; Software engineering; Data science; Software; Data mining; Software quality; Software development; Programming language; Set (abstract data type)","score_opus":0.05036001561640007,"score_gpt":0.2812554912041638,"score_spread":0.23089547558776374,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4292291355","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016974865,0.07811363,0.0074939174,0.88067746,0.006486571,0.000025905672,0.00031926282,0.00006223463,0.009846051],"genre_scores_gemma":[0.6161676,0.073521234,0.0137269655,0.25580147,0.03554881,0.00016659226,0.0006752111,0.00041265337,0.003979485],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9389174,0.034524657,0.0038015833,0.007523403,0.013848714,0.0013842678],"domain_scores_gemma":[0.35372305,0.52483314,0.021342464,0.016375424,0.07991143,0.003814475],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08336804,0.001614027,0.0029049418,0.00812288,0.0029399935,0.012334044,0.0074853655,0.014642189,0.007251221],"category_scores_gemma":[0.4054754,0.00074309926,0.0011164166,0.00910185,0.02100223,0.030441,0.0041377135,0.018526703,0.003474131],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009903355,0.00046556268,0.028105732,0.003816533,0.0009802658,0.00018065741,0.0069424827,0.0010858311,0.0009400254,0.31394812,0.18066372,0.4618807],"study_design_scores_gemma":[0.0004058098,0.00050684984,0.04932445,0.019637082,0.00094088493,0.000577128,0.021873537,0.0069463365,0.0020244198,0.64853215,0.24884984,0.0003814751],"about_ca_topic_score_codex":0.0092609925,"about_ca_topic_score_gemma":0.0088330535,"teacher_disagreement_score":0.91663194,"about_ca_system_score_codex":0.005062594,"about_ca_system_score_gemma":0.008644125,"threshold_uncertainty_score":0.44089758},"labels":[],"label_agreement":null},{"id":"W4292305354","doi":"10.2139/ssrn.4191846","title":"Self-Admitted Technical Debt in Commit Messages: Comparing Java, Python, and R","year":2022,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Commit; Java; Python (programming language); Computer science; Debt; Operating system; Programming language; World Wide Web; Business; Database; Finance","score_opus":0.009014154432849267,"score_gpt":0.24547349375658245,"score_spread":0.2364593393237332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4292305354","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9815746,0.00057498063,0.005418634,0.00063722144,0.000085985266,0.000048479906,0.0004721275,0.0013737263,0.009814206],"genre_scores_gemma":[0.99523723,0.00010705619,0.0024437085,0.00014495314,0.00002596627,0.00001789476,0.0004882777,0.00024193776,0.0012929288],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.990276,0.004069463,0.0006329451,0.0009417516,0.0031916797,0.00088815717],"domain_scores_gemma":[0.8421783,0.11509399,0.01479292,0.0125571685,0.011981277,0.003396405],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013006826,0.00033084716,0.00051747216,0.0020046998,0.000953697,0.0033139202,0.001780982,0.0013231522,0.0026983889],"category_scores_gemma":[0.11766945,0.00031499096,0.000540794,0.0021814643,0.0015502329,0.0047034184,0.0025507875,0.0019599174,0.00081094616],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.015127181,0.0029064578,0.53330344,0.0016238098,0.00067935407,0.0010023236,0.011203668,0.04634895,0.014896314,0.032847386,0.021566028,0.31849518],"study_design_scores_gemma":[0.0008275996,0.007331952,0.599565,0.00066684117,0.0009460534,0.0014646552,0.021022417,0.27810475,0.02288126,0.04219951,0.02456178,0.00042813987],"about_ca_topic_score_codex":0.009768583,"about_ca_topic_score_gemma":0.0098959375,"teacher_disagreement_score":0.9869932,"about_ca_system_score_codex":0.001412216,"about_ca_system_score_gemma":0.0027564047,"threshold_uncertainty_score":0.068787456},"labels":[],"label_agreement":null},{"id":"W4293080012","doi":"10.1007/s10515-022-00358-6","title":"Self-admitted technical debt in R: detection and causes","year":2022,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"University of British Columbia; Australian National University; University of Saskatchewan","keywords":"Computer science; Artificial intelligence; Source code; Domain (mathematical analysis); Software; Machine learning; Software quality; Semantics (computer science); Code (set theory); Software development; Data mining; Software engineering; Data science; Programming language; Set (abstract data type)","score_opus":0.0067949306597100394,"score_gpt":0.22774739982227968,"score_spread":0.22095246916256964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293080012","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.87908006,0.0032350246,0.09001971,0.0027148635,0.00020945244,0.0002683607,0.005381354,0.012211068,0.00688016],"genre_scores_gemma":[0.95177716,0.00053252425,0.039004363,0.00052447576,0.00010147462,0.00013947784,0.004317875,0.001238362,0.0023643493],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98413557,0.0044104313,0.0020375634,0.0035135434,0.005404041,0.0004989154],"domain_scores_gemma":[0.80095035,0.10443056,0.055103,0.015597915,0.021907477,0.0020106938],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.00937154,0.0007946255,0.0005677214,0.0069774636,0.0009202386,0.0020021372,0.0013184074,0.0010463798,0.0017479258],"category_scores_gemma":[0.11840104,0.00065104576,0.0007802916,0.005173792,0.0012097872,0.002548863,0.0022172723,0.0012529991,0.0011969552],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034555298,0.000112865404,0.811613,0.0012381752,0.00018612496,0.0024717625,0.004111466,0.004990362,0.005519328,0.004499961,0.016971992,0.14793938],"study_design_scores_gemma":[0.00008805212,0.00031056345,0.5952272,0.0017527671,0.00048576918,0.007953109,0.004528152,0.2710594,0.027937997,0.016688371,0.073606834,0.00036185296],"about_ca_topic_score_codex":0.005698003,"about_ca_topic_score_gemma":0.008823807,"teacher_disagreement_score":0.9906285,"about_ca_system_score_codex":0.0014617529,"about_ca_system_score_gemma":0.001910796,"threshold_uncertainty_score":0.049561977},"labels":[],"label_agreement":null},{"id":"W4293228199","doi":"10.1145/3511430.3511439","title":"Reproducibility Challenges and Their Impacts on Technical Q&amp;A Websites: The Practitioners’ Perspectives","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Viewpoints; Code review; Computer science; Code (set theory); Code smell; Reproducibility; Perspective (graphical); Data science; Software; Software quality; Software development; Programming language; Artificial intelligence","score_opus":0.061100026981090336,"score_gpt":0.31030737260071284,"score_spread":0.2492073456196225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293228199","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8193405,0.0058588185,0.022128172,0.09276482,0.00060726586,0.00032446824,0.00026808985,0.00032490856,0.05838301],"genre_scores_gemma":[0.9883265,0.0011547076,0.003666477,0.0040143393,0.00019070374,0.00013152164,0.00007996195,0.00010622241,0.0023294503],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.77090037,0.1430132,0.010272188,0.009842505,0.057403006,0.008568684],"domain_scores_gemma":[0.42174044,0.40439776,0.050941363,0.018463435,0.085600704,0.018856352],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.15395409,0.00071723614,0.0008741115,0.0048131067,0.010805505,0.022285884,0.002895996,0.005157306,0.004693204],"category_scores_gemma":[0.2870558,0.0011861393,0.0010053344,0.0045531495,0.0121014565,0.016337661,0.013409506,0.0063511096,0.0008891048],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001445096,0.00021700886,0.10569132,0.0012392311,0.00011004471,0.0020326923,0.77365696,0.000398763,0.0020509553,0.019667994,0.010672662,0.08411793],"study_design_scores_gemma":[0.0000406903,0.00035379716,0.05446357,0.0015692342,0.00009535343,0.0016312879,0.8255822,0.0017038784,0.0017499947,0.011408608,0.10117236,0.00022891953],"about_ca_topic_score_codex":0.0077396203,"about_ca_topic_score_gemma":0.0054186597,"teacher_disagreement_score":0.8460459,"about_ca_system_score_codex":0.009978538,"about_ca_system_score_gemma":0.014077949,"threshold_uncertainty_score":0.81419677},"labels":[],"label_agreement":null},{"id":"W4293228274","doi":"10.1145/3511430.3511463","title":"Commit-Checker: A human-centric approach for adopting bug inducing commit detection using machine learning models","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Commit; Citation; Computer science; Software; Software engineering; World Wide Web; Operating system; Database","score_opus":0.09667937685195939,"score_gpt":0.28727790948042364,"score_spread":0.19059853262846427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293228274","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008709622,0.0001823754,0.91356623,0.0013934684,0.00031611664,0.0005833172,0.00094876735,0.07070026,0.0035998425],"genre_scores_gemma":[0.18172246,0.00022717705,0.7991934,0.0009589957,0.00023394269,0.0005472266,0.0038076486,0.0048652617,0.008443949],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98775846,0.005070284,0.00083698786,0.002008319,0.003828452,0.0004975244],"domain_scores_gemma":[0.9489573,0.023701148,0.0029900754,0.0138508985,0.008903379,0.0015972229],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016647724,0.0019210532,0.0015210218,0.004043178,0.0014733536,0.0048636054,0.0055991276,0.0025036286,0.012845734],"category_scores_gemma":[0.058947913,0.0013309745,0.001968613,0.0017036857,0.0014283444,0.006014048,0.0046479744,0.0039043394,0.0055609625],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010655671,0.0014051984,0.016204204,0.0006830945,0.00046923553,0.0004886076,0.0006854915,0.11016993,0.012685132,0.019227074,0.1413912,0.6955253],"study_design_scores_gemma":[0.00007680801,0.00021916426,0.00088756153,0.00006466927,0.00005901871,0.00010395782,0.00011351705,0.9637782,0.008995119,0.01385949,0.011780189,0.000062297026],"about_ca_topic_score_codex":0.0062052286,"about_ca_topic_score_gemma":0.012458404,"teacher_disagreement_score":0.016647724,"about_ca_system_score_codex":0.0016423161,"about_ca_system_score_gemma":0.006194895,"threshold_uncertainty_score":0.08804262},"labels":[],"label_agreement":null},{"id":"W4293228314","doi":"10.1145/3511430.3511444","title":"Feature Transformation for Improved Software Bug Detection Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Software regression; Software bug; Machine learning; Artificial intelligence; Software; Classifier (UML); Random forest; Transformation (genetics); Precision and recall; Data mining; Feature selection; Feature (linguistics); Predictive modelling; Software quality; Data transformation; Software development; Programming language","score_opus":0.016094127414483374,"score_gpt":0.2398867386999696,"score_spread":0.22379261128548622,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293228314","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.112562746,0.00035110387,0.87917316,0.000318875,0.0000628428,0.000103987026,0.0005726169,0.0061384747,0.0007161699],"genre_scores_gemma":[0.780372,0.00015765357,0.21576712,0.00009499738,0.000044087607,0.00018206137,0.002153433,0.00023126705,0.0009973054],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988166,0.00034375876,0.00010325408,0.0003189903,0.00031060402,0.00010677726],"domain_scores_gemma":[0.99622726,0.002056722,0.00033343548,0.00042338984,0.00089790847,0.00006125636],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001985799,0.0010294395,0.0011085374,0.002132852,0.0004174755,0.00086493994,0.0014174127,0.0009076418,0.0011007698],"category_scores_gemma":[0.009296235,0.00041590925,0.0014660843,0.0016116757,0.0003728829,0.0014483649,0.0006981559,0.001756635,0.00089068833],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002203881,0.00030448774,0.012087524,0.00006422828,0.00011521061,0.00013275161,0.000086488966,0.67927945,0.0039525963,0.0018957026,0.0023026438,0.29955852],"study_design_scores_gemma":[0.000004707338,0.00002081885,0.00043365412,0.0000029972953,0.0000071510226,0.000018354041,0.0000048150923,0.9976624,0.00048382225,0.0011931126,0.0001639714,0.000004238865],"about_ca_topic_score_codex":0.010001552,"about_ca_topic_score_gemma":0.0069364468,"teacher_disagreement_score":0.010001552,"about_ca_system_score_codex":0.00084504083,"about_ca_system_score_gemma":0.0010394928,"threshold_uncertainty_score":0.019886672},"labels":[],"label_agreement":null},{"id":"W4293451258","doi":"10.1007/978-3-031-08530-7_57","title":"Detecting Use Case Scenarios in Requirements Artifacts: A Deep Learning Approach","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Classifier (UML); Transformer; Use Case Diagram; Artificial intelligence; Machine learning; Process (computing); Sequence diagram; Natural language; Requirements elicitation; Unified Modeling Language; Requirements analysis; Programming language; Class diagram; Software","score_opus":0.04203983792780778,"score_gpt":0.27241968032165065,"score_spread":0.23037984239384288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293451258","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25992668,0.00079303974,0.72478,0.00076812285,0.000058290265,0.00034730774,0.0024112614,0.0047440366,0.006171354],"genre_scores_gemma":[0.77488613,0.00031128534,0.21754411,0.00013589585,0.000025875674,0.00011310009,0.0038602417,0.00009703263,0.0030263215],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990103,0.00018277942,0.00008000042,0.00029275526,0.0002748131,0.00015916978],"domain_scores_gemma":[0.9969651,0.00171804,0.00045636785,0.0003242475,0.0003818928,0.00015432462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010308867,0.0010893357,0.00052302436,0.0029651935,0.0003643642,0.0016501294,0.0014917799,0.0012787648,0.0016237283],"category_scores_gemma":[0.0035855016,0.0005767159,0.0011322282,0.001479766,0.0003662113,0.0018635119,0.0014145707,0.0016896091,0.00081986137],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035061827,0.0012106653,0.043199584,0.00039845798,0.00022891819,0.00084763364,0.00046432676,0.13194673,0.01784763,0.004801222,0.009170439,0.7895337],"study_design_scores_gemma":[0.0000050960716,0.000041870902,0.0031926115,0.000040323608,0.000025242221,0.00011555305,0.00010137084,0.9865753,0.0033316605,0.0054939203,0.0010668775,0.000010134342],"about_ca_topic_score_codex":0.006665707,"about_ca_topic_score_gemma":0.014264398,"teacher_disagreement_score":0.006665707,"about_ca_system_score_codex":0.00100891,"about_ca_system_score_gemma":0.0009987667,"threshold_uncertainty_score":0.013253808},"labels":[],"label_agreement":null},{"id":"W4294175720","doi":"10.5753/sbes.2008.21335","title":"Uso de Gerência de Conhecimento para Apoiar a Rastreabilidade e a Avaliação de Impacto de Alterações","year":2008,"lang":"pt","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Impact","funders":"","keywords":"Humanities; Philosophy","score_opus":0.04723508797699636,"score_gpt":0.3171886744728873,"score_spread":0.2699535864958909,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294175720","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81392217,0.0027306818,0.15062821,0.0012137209,0.00017393517,0.00027518024,0.00020958229,0.0032824562,0.027563987],"genre_scores_gemma":[0.9602673,0.0004276595,0.035924114,0.00008717555,0.000019560242,0.000054743912,0.000069504524,0.00017758505,0.0029724408],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99619603,0.0008878768,0.0001886399,0.0007971607,0.0016903486,0.00023993889],"domain_scores_gemma":[0.9882208,0.0051839957,0.0013640475,0.0027216997,0.0020102318,0.00049921725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030736802,0.0007796027,0.00074671075,0.0012584182,0.0008253668,0.003192691,0.0014768161,0.0014461895,0.0046546813],"category_scores_gemma":[0.020269949,0.00048623292,0.0005721369,0.0010900149,0.0015288347,0.003933578,0.0018084946,0.0008643628,0.00061471475],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001618306,0.0009336456,0.042657666,0.0016452824,0.0002632296,0.0007580069,0.008387979,0.04714448,0.16865736,0.03216151,0.0033509035,0.6924217],"study_design_scores_gemma":[0.00044751985,0.0041965498,0.22620185,0.00080961495,0.0011687453,0.001894023,0.009561113,0.35430485,0.22031835,0.09024978,0.09018698,0.0006606338],"about_ca_topic_score_codex":0.005143228,"about_ca_topic_score_gemma":0.0053633875,"teacher_disagreement_score":0.005143228,"about_ca_system_score_codex":0.0013629905,"about_ca_system_score_gemma":0.001576356,"threshold_uncertainty_score":0.016255379},"labels":[],"label_agreement":null},{"id":"W4294237828","doi":"10.1002/smr.2505","title":"Improving the detection of community smells through socio‐technical and sentiment analysis","year":2022,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code smell; Computer science; Classifier (UML); Software development; Software; Benchmark (surveying); Data science; Empirical research; Software quality; Knowledge management; Artificial intelligence; Software engineering; Machine learning","score_opus":0.016791062457961817,"score_gpt":0.2797237609382808,"score_spread":0.26293269848031897,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294237828","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8223284,0.0009261981,0.16801205,0.00101098,0.00019851525,0.0002211478,0.00097289553,0.0020112488,0.0043186187],"genre_scores_gemma":[0.9652754,0.00012826953,0.032353707,0.00009069815,0.000087332104,0.000049678252,0.00090611266,0.000056003362,0.0010529256],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99872786,0.00043106024,0.00007954438,0.0002558633,0.00037968173,0.0001259691],"domain_scores_gemma":[0.994445,0.00201886,0.0011743187,0.00027208307,0.0017766358,0.00031302293],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002002868,0.0008885708,0.0006980285,0.003091693,0.00042968336,0.000892551,0.00060457597,0.00091381883,0.000776239],"category_scores_gemma":[0.0056423177,0.0001808892,0.0006463237,0.00095268263,0.00026671248,0.0012539887,0.0010020809,0.00087281613,0.0006920663],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008059139,0.001555427,0.25333485,0.0006936472,0.00048765328,0.0007778007,0.0010617645,0.03803019,0.090669625,0.0012974642,0.017643435,0.59364223],"study_design_scores_gemma":[0.000018598592,0.00018168667,0.053581458,0.000039001185,0.00006161125,0.000093992276,0.0004083179,0.934127,0.008229949,0.0012878097,0.0019348038,0.000035877976],"about_ca_topic_score_codex":0.0029696983,"about_ca_topic_score_gemma":0.003757875,"teacher_disagreement_score":0.003091693,"about_ca_system_score_codex":0.0004912145,"about_ca_system_score_gemma":0.0004399698,"threshold_uncertainty_score":0.010592282},"labels":[],"label_agreement":null},{"id":"W4294529940","doi":"10.1145/3551349.3559547","title":"End-to-End Rationale Reconstruction","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Correctness; Exploit; Traceability; Pipeline (software); Information extraction; Kernel (algebra); Program comprehension; Artificial intelligence; Context (archaeology); Machine learning; Process (computing); Software engineering; Data science; Programming language","score_opus":0.031809839870279855,"score_gpt":0.28056785322802186,"score_spread":0.248758013357742,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294529940","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0057362057,0.00024724522,0.97657996,0.0006430948,0.000082554274,0.00055083615,0.0022970627,0.0084396815,0.005423379],"genre_scores_gemma":[0.070402175,0.00041090616,0.9155399,0.0002849322,0.00003609478,0.00037231744,0.0069811526,0.0014869431,0.0044855378],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9939004,0.0016251976,0.00045936316,0.0009354712,0.0027650078,0.00031452382],"domain_scores_gemma":[0.9828212,0.008003053,0.000955154,0.0040177912,0.003919905,0.00028303397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005912215,0.001951381,0.0009285233,0.0047763963,0.0009882391,0.004912531,0.0026251688,0.0021642086,0.013840454],"category_scores_gemma":[0.03357523,0.001106943,0.002373656,0.0018633921,0.0011745929,0.0045305924,0.004464184,0.0038090872,0.00765533],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035003177,0.00045272429,0.004583788,0.0016232429,0.00018165178,0.0015858756,0.0019484943,0.03192801,0.025211863,0.09457638,0.0331586,0.8043994],"study_design_scores_gemma":[0.00015013026,0.00026491872,0.003235929,0.0010027022,0.0002611846,0.0013217619,0.0018806137,0.4234599,0.0785459,0.29278004,0.1969035,0.00019350302],"about_ca_topic_score_codex":0.0026871203,"about_ca_topic_score_gemma":0.0052099316,"teacher_disagreement_score":0.013840454,"about_ca_system_score_codex":0.0013848565,"about_ca_system_score_gemma":0.004580619,"threshold_uncertainty_score":0.046300948},"labels":[],"label_agreement":null},{"id":"W4295277053","doi":"10.1016/j.jss.2022.111505","title":"A survey of software architectural change detection and categorization techniques","year":2022,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Categorization; Software engineering; Computer science; Change detection; Software; Architectural pattern; Data science; Systems engineering; Software development; Engineering; Artificial intelligence; Software design; Programming language","score_opus":0.031604897148514915,"score_gpt":0.2580734554457964,"score_spread":0.2264685582972815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4295277053","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.082465865,0.22640191,0.6662307,0.0028233957,0.0006833107,0.0006725585,0.002441951,0.007278117,0.0110021755],"genre_scores_gemma":[0.22946359,0.07772559,0.6767752,0.0010959257,0.0009383094,0.00039555854,0.0065339534,0.0005682809,0.0065035596],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99556905,0.0005894662,0.00057634985,0.001276873,0.0017536756,0.00023463741],"domain_scores_gemma":[0.9907073,0.004179767,0.00087578007,0.0012528555,0.0027491732,0.00023518714],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033668398,0.0013459054,0.0024297857,0.014278646,0.0011812415,0.0030607062,0.003477831,0.0015448859,0.0014200548],"category_scores_gemma":[0.009274846,0.00077740935,0.0019530051,0.013137267,0.0006879262,0.0051385066,0.001085406,0.0013291714,0.0012058091],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008947183,0.00028167112,0.009140281,0.0014638731,0.00010759252,0.00005474182,0.00019968025,0.0017726349,0.004382231,0.0024412163,0.006102392,0.9739644],"study_design_scores_gemma":[0.00017546855,0.0020188536,0.10119832,0.00444766,0.0016702745,0.006153009,0.0034047593,0.42251647,0.07695058,0.082177356,0.29860032,0.0006869233],"about_ca_topic_score_codex":0.00493542,"about_ca_topic_score_gemma":0.006056388,"teacher_disagreement_score":0.014278646,"about_ca_system_score_codex":0.00097205344,"about_ca_system_score_gemma":0.0020173073,"threshold_uncertainty_score":0.017805755},"labels":[],"label_agreement":null},{"id":"W4296087229","doi":"10.3390/app12189017","title":"Evidence-Based Software Engineering: A Checklist-Based Approach to Assess the Abstracts of Reviews Self-Identifying as Systematic Reviews","year":2022,"lang":"en","type":"article","venue":"Applied Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Systematic review; Checklist; Computer science; Management science; Data science; Engineering ethics; Psychology; Risk analysis (engineering); MEDLINE; Medicine; Political science; Engineering; Cognitive psychology","score_opus":0.16820838240296807,"score_gpt":0.3307259528974853,"score_spread":0.16251757049451723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296087229","genre_codex":"protocol","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037168479,0.030132659,0.40355223,0.012036916,0.0034863683,0.5193005,0.012721291,0.0037380042,0.01131515],"genre_scores_gemma":[0.0055111395,0.0051835575,0.8341346,0.00085886015,0.00017024716,0.15156718,0.0017880391,0.00019939097,0.00058702653],"study_design_codex":"systematic_review","study_design_gemma":"observational","domain_scores_codex":[0.2802054,0.43293333,0.2188044,0.011189134,0.054353844,0.0025139763],"domain_scores_gemma":[0.22935273,0.4289213,0.0936561,0.040309895,0.20009984,0.007660197],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.47638154,0.010994661,0.017497035,0.09168075,0.007815364,0.017918259,0.013426941,0.009153719,0.012788219],"category_scores_gemma":[0.6578933,0.0076072807,0.027873052,0.0668128,0.010789743,0.01639148,0.020647261,0.013479952,0.0062910016],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014911871,0.00048643947,0.0028412133,0.48431674,0.008934155,0.0007321787,0.016581053,0.002727203,0.0042910227,0.033447426,0.06423462,0.37991676],"study_design_scores_gemma":[0.005128226,0.0018658473,0.0108741205,0.4092826,0.017897781,0.001910989,0.011137934,0.009994393,0.0074811894,0.10140041,0.4207524,0.0022740536],"about_ca_topic_score_codex":0.0066981665,"about_ca_topic_score_gemma":0.013385234,"teacher_disagreement_score":0.52361846,"about_ca_system_score_codex":0.02406823,"about_ca_system_score_gemma":0.11164317,"threshold_uncertainty_score":0.6457148},"labels":[],"label_agreement":null},{"id":"W4296132134","doi":"10.1145/3563214","title":"Video Game Bad Smells: What They Are and How Developers Perceive Them","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Game Developer; Video game development; Video game; Code smell; Relevance (law); Game design; Game development tool; Game testing; Animation; World Wide Web; Software; Game art design; Software development; Game design document; Multimedia; Data science; Software quality","score_opus":0.08737181515015419,"score_gpt":0.29583185583088095,"score_spread":0.20846004068072677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296132134","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9702262,0.0035343163,0.013462633,0.0020156468,0.00014180123,0.00017702904,0.00025047944,0.0005213987,0.00967046],"genre_scores_gemma":[0.9872953,0.0012153015,0.0066795773,0.00059345487,0.000054095697,0.00009349198,0.00037876246,0.00022210412,0.0034678166],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9912799,0.0033861147,0.000712414,0.00067348283,0.0032986486,0.0006494749],"domain_scores_gemma":[0.958,0.02350099,0.0096844975,0.0016080064,0.005489659,0.001716777],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061164643,0.00077374204,0.00056011707,0.0034213152,0.0012236696,0.0036323573,0.0007179268,0.0016976222,0.0011953304],"category_scores_gemma":[0.059260238,0.00052094215,0.00040135437,0.0016561978,0.001564093,0.0045250645,0.0029956077,0.0012160856,0.0004547568],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073863164,0.00025747583,0.4279123,0.002121306,0.00014069115,0.0050270776,0.27440208,0.0010327891,0.032361854,0.0052368836,0.015784107,0.2349848],"study_design_scores_gemma":[0.000052152027,0.000757522,0.5259196,0.002658885,0.0001992089,0.011237011,0.31297794,0.00927815,0.0093540605,0.008792397,0.11836229,0.0004107481],"about_ca_topic_score_codex":0.00306757,"about_ca_topic_score_gemma":0.0044915094,"teacher_disagreement_score":0.0061164643,"about_ca_system_score_codex":0.0012871069,"about_ca_system_score_gemma":0.0007787681,"threshold_uncertainty_score":0.03234732},"labels":[],"label_agreement":null},{"id":"W4296422560","doi":"10.1145/3544902.3546639","title":"Example Driven Code Review Explanation","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Code review; Computer science; Software inspection; Software quality; Source code; Word embedding; Code (set theory); Software; Context (archaeology); Empirical research; Set (abstract data type); Software engineering; Software development; Data mining; Information retrieval; Artificial intelligence; Embedding; Programming language","score_opus":0.04780188340260429,"score_gpt":0.2913197945639388,"score_spread":0.24351791116133453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296422560","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06102685,0.002429464,0.8900487,0.0050704223,0.0006528626,0.002896795,0.00340114,0.019987246,0.014486613],"genre_scores_gemma":[0.1970494,0.0010468605,0.78633547,0.0009363494,0.00025693743,0.0013441159,0.003353863,0.00081401464,0.008862976],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9880183,0.0063999128,0.0011015618,0.001110481,0.0030850843,0.00028470243],"domain_scores_gemma":[0.8811087,0.06967753,0.010192819,0.008790611,0.029253088,0.0009772045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008018056,0.0012430548,0.0006118483,0.0028915766,0.00067661103,0.0017451263,0.0018047394,0.0015213346,0.010894227],"category_scores_gemma":[0.08044587,0.00041323024,0.0007684467,0.0016490191,0.00047988901,0.0018779321,0.0014534628,0.000939259,0.004420325],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008184158,0.00022870791,0.010375189,0.003992984,0.00014669196,0.00078293553,0.0030706122,0.01017793,0.020200532,0.0071233846,0.063058704,0.8800239],"study_design_scores_gemma":[0.00062824675,0.0016027994,0.024750149,0.0038130372,0.00049475057,0.0044951187,0.0032158515,0.367096,0.08175582,0.038025882,0.4735443,0.00057806465],"about_ca_topic_score_codex":0.0016718825,"about_ca_topic_score_gemma":0.0032567699,"teacher_disagreement_score":0.010894227,"about_ca_system_score_codex":0.0009916136,"about_ca_system_score_gemma":0.0027466777,"threshold_uncertainty_score":0.042404056},"labels":[],"label_agreement":null},{"id":"W4297034603","doi":"","title":"The CoLiS Platform for the Analysis of Maintainer Scripts in Debian Software Packages","year":2022,"lang":"en","type":"article","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure","funders":"","keywords":"Software maintainer; Computer science; Scripting language; Software; Software engineering; Operating system","score_opus":0.015633073568107756,"score_gpt":0.23969103949363663,"score_spread":0.22405796592552887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297034603","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009347727,0.00034871508,0.3013972,0.00015793307,0.00017041073,0.00042328332,0.013586123,0.66886926,0.00569934],"genre_scores_gemma":[0.08502142,0.00073872483,0.5418296,0.00054443104,0.00017759875,0.0016027268,0.10071668,0.25829372,0.011075202],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9958703,0.0007787641,0.0005101323,0.0010635863,0.0012828804,0.0004944023],"domain_scores_gemma":[0.9925006,0.0030502859,0.00071327057,0.0021620495,0.0011795303,0.00039428915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050966935,0.003853174,0.0020613659,0.0055545955,0.0015360175,0.0048649167,0.004692017,0.0013695295,0.039455097],"category_scores_gemma":[0.016785678,0.0024565882,0.00403774,0.0045023533,0.0013146331,0.004887497,0.004324688,0.0038144249,0.022839816],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029821522,0.0010287766,0.012070078,0.00370479,0.0013108711,0.0010956589,0.0027195122,0.015823642,0.037931774,0.04116532,0.47133482,0.4088325],"study_design_scores_gemma":[0.0015174309,0.00052137166,0.01827628,0.0013524811,0.0006342125,0.0012328621,0.0009427515,0.32095695,0.09444122,0.05590757,0.5034301,0.00078669685],"about_ca_topic_score_codex":0.008464272,"about_ca_topic_score_gemma":0.008748374,"teacher_disagreement_score":0.039455097,"about_ca_system_score_codex":0.0016983261,"about_ca_system_score_gemma":0.004774563,"threshold_uncertainty_score":0.13199043},"labels":[],"label_agreement":null},{"id":"W4297421247","doi":"","title":"Recovering Binary Class Relationships: Putting Icing on the UML Cake","year":2004,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Unified Modeling Language; Computer science; Class (philosophy); Binary number; Class diagram; Programming language; Icing; Theoretical computer science; Software engineering; Artificial intelligence; Mathematics; Arithmetic; Meteorology; Software; Physics","score_opus":0.07254448393617577,"score_gpt":0.2850188166278928,"score_spread":0.21247433269171703,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297421247","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018266149,0.0006162352,0.9715887,0.0036012463,0.00025844813,0.00005666317,0.00008415553,0.0020068705,0.0035215546],"genre_scores_gemma":[0.19390833,0.00074841897,0.79707676,0.0014338085,0.0002899804,0.000114656825,0.00024349707,0.0018284239,0.0043560807],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98704964,0.0046873917,0.0006785102,0.001672225,0.0052242978,0.0006879865],"domain_scores_gemma":[0.9494938,0.022027899,0.004436752,0.019181762,0.0037837469,0.0010759537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01188833,0.0010773086,0.001401525,0.0059366496,0.0029473144,0.0098676095,0.0029349935,0.004551339,0.0044311057],"category_scores_gemma":[0.08675698,0.0017825731,0.0013331205,0.0034929498,0.009652081,0.022364631,0.010215518,0.008291787,0.0016114563],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031252208,0.00010039317,0.0061407695,0.000256738,0.000054032484,0.0005166644,0.0046208682,0.012374309,0.0107222805,0.46736336,0.01153509,0.48600292],"study_design_scores_gemma":[0.00007173805,0.00007057852,0.0017492066,0.00043741992,0.00008583816,0.0007496687,0.0018558106,0.16137996,0.01940495,0.7132059,0.100806415,0.00018252547],"about_ca_topic_score_codex":0.006654937,"about_ca_topic_score_gemma":0.005913384,"teacher_disagreement_score":0.01188833,"about_ca_system_score_codex":0.0023227779,"about_ca_system_score_gemma":0.00311667,"threshold_uncertainty_score":0.06287223},"labels":[],"label_agreement":null},{"id":"W4297747149","doi":"","title":"Impact of switching bug trackers: a case study on a medium-sized open source project","year":2019,"lang":"en","type":"preprint","venue":"INRIA a CCSD electronic archive server","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure","funders":"","keywords":"Open source; BitTorrent tracker; Open source software; Computer science; Software bug; Software; Programming language; Artificial intelligence; Eye tracking","score_opus":0.03511774202131468,"score_gpt":0.3539059257747847,"score_spread":0.31878818375347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297747149","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9950368,0.0002575534,0.0016142407,0.00044746723,0.000023196832,0.00007349881,0.00011697547,0.00030332507,0.0021269391],"genre_scores_gemma":[0.99574924,0.00012301994,0.002775498,0.000085505984,0.00001396887,0.00003257569,0.00018799293,0.00009926827,0.00093287136],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9907945,0.003938981,0.0005327669,0.001071833,0.002799548,0.0008623227],"domain_scores_gemma":[0.8801201,0.08579013,0.011355305,0.0074609662,0.008222324,0.0070510716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008516379,0.0007719219,0.0004680224,0.0031446638,0.0020439082,0.0023365626,0.0018888948,0.002476572,0.0025907168],"category_scores_gemma":[0.052463856,0.00050512154,0.0007110267,0.0029843324,0.0018083225,0.0025979795,0.002533295,0.0022168083,0.00045872963],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033589113,0.011460832,0.48614073,0.0014474423,0.000632683,0.020473773,0.03397225,0.03060416,0.027290195,0.0039968058,0.011602492,0.36901975],"study_design_scores_gemma":[0.0012309288,0.015736831,0.7523076,0.0010234762,0.0017341604,0.011861231,0.057764977,0.09517328,0.019234607,0.008487688,0.03498982,0.0004553301],"about_ca_topic_score_codex":0.008397192,"about_ca_topic_score_gemma":0.011463171,"teacher_disagreement_score":0.008516379,"about_ca_system_score_codex":0.0014020744,"about_ca_system_score_gemma":0.0019441943,"threshold_uncertainty_score":0.045039475},"labels":[],"label_agreement":null},{"id":"W4297924730","doi":"10.1007/978-3-031-16078-3_49","title":"A Deep Learning Approach to UML Class Diagrams Discovery from Textual Specifications of Software Systems","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Rimouski","funders":"","keywords":"Computer science; Class diagram; Programming language; Unified Modeling Language; Artificial intelligence; Class (philosophy); Natural language processing; Natural language; Software engineering; Software","score_opus":0.026479043178970603,"score_gpt":0.22553957714422843,"score_spread":0.19906053396525783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297924730","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017293878,0.00057858514,0.9751702,0.0005342171,0.00003810809,0.00008270976,0.0008275099,0.004042439,0.0014323536],"genre_scores_gemma":[0.17289384,0.0004682545,0.81590825,0.00034666157,0.00006186229,0.00016422269,0.0044297865,0.0003165063,0.0054106894],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99899334,0.00026586873,0.00010612171,0.00026327145,0.00027612929,0.00009520303],"domain_scores_gemma":[0.99522424,0.0034504477,0.00020062135,0.00046283257,0.00055725756,0.000104561506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016132533,0.00075900886,0.00085161184,0.002319847,0.00062971143,0.0016294383,0.0025501628,0.0014706799,0.0029219226],"category_scores_gemma":[0.005137732,0.0008901652,0.0015327046,0.0018117246,0.00059307105,0.0022649826,0.0019438622,0.0023401526,0.0010142124],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001495813,0.00029011478,0.0019217953,0.00025591505,0.00012593523,0.00014148507,0.00024688663,0.111586004,0.0050800303,0.015390185,0.01006654,0.8547454],"study_design_scores_gemma":[0.000008968588,0.000017834553,0.00026573366,0.000029358083,0.000017453178,0.000025364598,0.000029022387,0.97744924,0.0014794695,0.018773085,0.001897823,0.0000066970197],"about_ca_topic_score_codex":0.01619835,"about_ca_topic_score_gemma":0.032329053,"teacher_disagreement_score":0.01619835,"about_ca_system_score_codex":0.001410362,"about_ca_system_score_gemma":0.0015462538,"threshold_uncertainty_score":0.032208085},"labels":[],"label_agreement":null},{"id":"W4300546131","doi":"10.7287/peerj.preprints.2617","title":"Curating GitHub for engineered software projects","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Classifier (UML); Software engineering; Set (abstract data type); Software metric; Software development; Software bug; Skew; Data mining; Machine learning; Software quality; Artificial intelligence; Programming language","score_opus":0.042283290895837815,"score_gpt":0.2967354350366486,"score_spread":0.2544521441408108,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4300546131","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.72963697,0.00470881,0.1309038,0.0023647437,0.00049387,0.002762731,0.07653728,0.034423314,0.018168464],"genre_scores_gemma":[0.47917128,0.0017380932,0.36834854,0.00038651653,0.00016807744,0.0018139834,0.13869962,0.0027562606,0.0069176345],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9912901,0.0014314743,0.0013589048,0.0020822673,0.0033109838,0.0005262309],"domain_scores_gemma":[0.9504173,0.011202138,0.0144771505,0.011353041,0.010683198,0.0018672642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077548143,0.0013460367,0.0007814755,0.024407767,0.001587618,0.0031453301,0.0017727561,0.0011525155,0.002047325],"category_scores_gemma":[0.053814694,0.0006623246,0.0009609798,0.016224753,0.0011447961,0.0033185773,0.006240573,0.0014492122,0.0023525637],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047560164,0.00025150986,0.4143955,0.003244899,0.00030990606,0.0021701285,0.0104253935,0.003960328,0.01234407,0.008884588,0.11496254,0.4285756],"study_design_scores_gemma":[0.00009770742,0.00029637266,0.63351446,0.0012292508,0.00022274241,0.0025649224,0.007176461,0.046892934,0.0298673,0.011289372,0.2665604,0.00028808464],"about_ca_topic_score_codex":0.010208414,"about_ca_topic_score_gemma":0.02341218,"teacher_disagreement_score":0.024407767,"about_ca_system_score_codex":0.0015071637,"about_ca_system_score_gemma":0.0036768927,"threshold_uncertainty_score":0.04101187},"labels":[],"label_agreement":null},{"id":"W4300564642","doi":"10.1145/3196398.3196434","title":"The Android update problem","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Android (operating system); Computer science; Phone; Android application; Merge (version control); Empirical research; World Wide Web; Operating system; Information retrieval","score_opus":0.017176460711947693,"score_gpt":0.2702848470659632,"score_spread":0.2531083863540155,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4300564642","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6476053,0.010903479,0.16668825,0.025258448,0.0020541437,0.001387776,0.0053745573,0.017481564,0.12324653],"genre_scores_gemma":[0.8931955,0.0028497698,0.061359163,0.0050526313,0.0009799729,0.00049266097,0.003832498,0.003480786,0.02875704],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9823439,0.0025600148,0.0015573368,0.0035785113,0.008963991,0.0009961888],"domain_scores_gemma":[0.9327853,0.03570319,0.008636846,0.013975495,0.007693441,0.0012056658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005206878,0.0010466973,0.00082450494,0.0033454879,0.0024303382,0.0032649236,0.0028236194,0.0040547075,0.0053374604],"category_scores_gemma":[0.074455395,0.0011654638,0.0010548347,0.0032390882,0.0024303058,0.00973895,0.004015546,0.0036075173,0.002637425],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091312244,0.0004047391,0.13435721,0.0014650266,0.00025257372,0.010396302,0.017518288,0.0072410814,0.011141621,0.06610466,0.0957742,0.6544312],"study_design_scores_gemma":[0.00020167838,0.00044746115,0.10579763,0.0011352467,0.0005353885,0.043385003,0.011665121,0.04977623,0.017850447,0.06051464,0.70832115,0.0003699578],"about_ca_topic_score_codex":0.0062050223,"about_ca_topic_score_gemma":0.0051265564,"teacher_disagreement_score":0.0062050223,"about_ca_system_score_codex":0.0015010351,"about_ca_system_score_gemma":0.0017500944,"threshold_uncertainty_score":0.027536929},"labels":[],"label_agreement":null},{"id":"W4300889023","doi":"10.1145/3133956","title":"Proceedings of the 2017 ACM SIGSAC Conference on Computer and Communications Security","year":2017,"lang":"en","type":"paratext","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":68,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Presentation (obstetrics); Computer science; Government (linguistics); Inclusion (mineral); Library science; Process (computing); Operations research; Sociology; Engineering; Medicine","score_opus":0.07594937142261006,"score_gpt":0.3353062954452068,"score_spread":0.25935692402259675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4300889023","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010454131,0.06336578,0.0406054,0.09526891,0.3636047,0.0021228348,0.008431509,0.005981096,0.41016567],"genre_scores_gemma":[0.033334676,0.047386304,0.017995946,0.013715723,0.050349884,0.0011976137,0.02097553,0.0027727748,0.8122715],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99405056,0.0014354577,0.00035781477,0.00070348737,0.0030585318,0.0003940791],"domain_scores_gemma":[0.9795436,0.0036204522,0.00060402195,0.0019034101,0.010176781,0.0041516805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075187143,0.0014437016,0.0016488433,0.002956459,0.0027811946,0.009344839,0.0018974737,0.0019868219,0.19011688],"category_scores_gemma":[0.01771755,0.00053275283,0.001218597,0.0015116013,0.0017569714,0.008005917,0.0037916824,0.005030917,0.15558192],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000072198985,0.00004889609,0.00035042007,0.00016898656,0.000019211766,0.000088804314,0.000109762375,0.0001110844,0.00045348742,0.0029873808,0.9396439,0.05594585],"study_design_scores_gemma":[0.00001063609,0.00002914482,0.00041620957,0.0002215417,0.000009760467,0.000104563,0.0001950363,0.00040335394,0.00026978817,0.0026472437,0.9956761,0.000016620477],"about_ca_topic_score_codex":0.0031326013,"about_ca_topic_score_gemma":0.004510752,"teacher_disagreement_score":0.19011688,"about_ca_system_score_codex":0.0027217593,"about_ca_system_score_gemma":0.006849497,"threshold_uncertainty_score":0.6360043},"labels":[],"label_agreement":null},{"id":"W4300960622","doi":"10.1007/s10664-022-10223-5","title":"A controlled experiment of different code representations for learning-based program repair","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Code (set theory); Source code; Representation (politics); Artificial intelligence; Point (geometry); Syntax; Programming language; Perspective (graphical); Process (computing); Natural language processing; Set (abstract data type)","score_opus":0.03017434424875222,"score_gpt":0.33110164950214477,"score_spread":0.30092730525339256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4300960622","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98787177,0.000117832235,0.006944164,0.00007006475,0.00013803548,0.0030472132,0.00038926088,0.00031660826,0.0011051494],"genre_scores_gemma":[0.96962625,0.00016871,0.020630563,0.00022597282,0.000074833915,0.0052091903,0.00064000697,0.0001345176,0.003289957],"study_design_codex":"nonrandomized_trial","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972651,0.00088994927,0.00044590642,0.0007578822,0.00038534252,0.00025585378],"domain_scores_gemma":[0.9516662,0.038479738,0.0025755863,0.003752738,0.001680864,0.0018448676],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036833775,0.0013417272,0.0012224143,0.00053519406,0.00059527735,0.0011023228,0.0022701393,0.0021366847,0.010339895],"category_scores_gemma":[0.03323672,0.0008761733,0.000648715,0.0003312537,0.0012499796,0.0018614321,0.0010859765,0.0019963915,0.001042833],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.29330128,0.34542656,0.005654008,0.0034596324,0.0005256359,0.00031221536,0.004958994,0.012162586,0.16482407,0.002297227,0.0028073965,0.16427042],"study_design_scores_gemma":[0.12228468,0.67823935,0.040410724,0.00040748215,0.0015525065,0.00031233297,0.0017983759,0.039444145,0.10240035,0.005772006,0.006854767,0.00052331627],"about_ca_topic_score_codex":0.0015294076,"about_ca_topic_score_gemma":0.0015119208,"teacher_disagreement_score":0.010339895,"about_ca_system_score_codex":0.0006946887,"about_ca_system_score_gemma":0.001199789,"threshold_uncertainty_score":0.034590364},"labels":[],"label_agreement":null},{"id":"W4301135585","doi":"10.7287/peerj.preprints.1138v1","title":"The charming code that error messages are talking about","year":2015,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Python (programming language); Cyclomatic complexity; Debugging; Computer science; Programming language; Random testing; Source lines of code; Code coverage; Software quality; Software bug; Syntax error; Software; Charm (quantum number); Software metric; Source code; Code (set theory); Abstract syntax tree; Software development; Test case; Particle physics; Machine learning","score_opus":0.07479187798440508,"score_gpt":0.328330902469312,"score_spread":0.25353902448490695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4301135585","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29709652,0.0061115306,0.41634974,0.027516391,0.009213745,0.0012314975,0.006745044,0.0630041,0.17273147],"genre_scores_gemma":[0.667053,0.002284665,0.18001258,0.012492163,0.0015998723,0.0007260799,0.0037412816,0.016190862,0.11589956],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99168605,0.002464841,0.00059205794,0.0008665065,0.0039257654,0.00046480598],"domain_scores_gemma":[0.950754,0.019820156,0.010392356,0.007806764,0.010174995,0.001051796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026141615,0.0015769986,0.00058862503,0.0024870678,0.0020595687,0.0028588956,0.000929551,0.0022115104,0.013503694],"category_scores_gemma":[0.046489693,0.00055378984,0.00048277265,0.002005445,0.0034545278,0.0053307414,0.0027597174,0.003051101,0.006618674],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015626523,0.0003344461,0.058407843,0.0029268577,0.00027436137,0.0040674014,0.01627849,0.0043713422,0.047951967,0.1461075,0.21124153,0.5064757],"study_design_scores_gemma":[0.00007202816,0.00040193554,0.02924158,0.0022151973,0.00022702881,0.0064423615,0.0035824536,0.010631853,0.06547484,0.05838846,0.8229501,0.00037212577],"about_ca_topic_score_codex":0.0021918914,"about_ca_topic_score_gemma":0.0018684452,"teacher_disagreement_score":0.013503694,"about_ca_system_score_codex":0.0012060588,"about_ca_system_score_gemma":0.0016376926,"threshold_uncertainty_score":0.04517436},"labels":[],"label_agreement":null},{"id":"W4301155764","doi":"","title":"Improving Semantic Transparency of Committee-Designed Languages through Crowd-sourcing","year":2014,"lang":"en","type":"article","venue":"Espace ÉTS (ETS)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Crowd sourcing; Transparency (behavior); Computer science; World Wide Web; Computer security","score_opus":0.013036716222754248,"score_gpt":0.26088629613466413,"score_spread":0.2478495799119099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4301155764","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15865147,0.00019155063,0.81905013,0.0015347868,0.0001265335,0.0010227795,0.00036170968,0.00421023,0.014850789],"genre_scores_gemma":[0.5189551,0.00014002701,0.4712167,0.00041295486,0.00004716495,0.0007978701,0.0009868913,0.0013008355,0.0061424393],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.96666235,0.022615794,0.0017460635,0.0030688907,0.004836684,0.0010702196],"domain_scores_gemma":[0.9079933,0.052328497,0.0045425366,0.021601846,0.010896843,0.0026369053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.039226193,0.0016483944,0.0010189979,0.0029252726,0.0025975248,0.00775196,0.0033870782,0.0029197852,0.003716073],"category_scores_gemma":[0.06967537,0.00089734234,0.0019685717,0.0015904383,0.0040663104,0.008634083,0.013035488,0.0033312219,0.0010971596],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027900168,0.0023939582,0.012833884,0.001962509,0.00039233247,0.0035073035,0.14270283,0.06794872,0.15862945,0.19051784,0.0150082065,0.40131292],"study_design_scores_gemma":[0.0007503451,0.0013216152,0.007119054,0.0011565362,0.00034781825,0.0012456734,0.0483122,0.37895915,0.10694061,0.25329077,0.19952159,0.0010346554],"about_ca_topic_score_codex":0.0032794804,"about_ca_topic_score_gemma":0.0031837577,"teacher_disagreement_score":0.039226193,"about_ca_system_score_codex":0.0028416382,"about_ca_system_score_gemma":0.0053143348,"threshold_uncertainty_score":0.20745039},"labels":[],"label_agreement":null},{"id":"W4301325890","doi":"10.29173/istl1596","title":"The Impact of Data Source on the Ranking of Computer Scientists Based on Citation Indicators: A Comparison of Web of Science and Scopus.","year":2014,"lang":"en","type":"article","venue":"Issues in Science and Technology Librarianship","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Scopus; Citation; Web of science; Ranking (information retrieval); Citation impact; Computer science; Citation analysis; Journal ranking; Impact factor; Bibliometrics; Information retrieval; Data science; World Wide Web; Library science; Political science; MEDLINE","score_opus":0.040395890902166144,"score_gpt":0.340836562214279,"score_spread":0.30044067131211283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4301325890","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8953223,0.040392876,0.006523016,0.009086421,0.000807072,0.00051023305,0.012712771,0.00034747767,0.034297794],"genre_scores_gemma":[0.98460937,0.0043640444,0.004882739,0.00042267403,0.0001985925,0.0001370681,0.0045478656,0.000098188604,0.0007394456],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94764346,0.020253459,0.0060611703,0.0016701092,0.023386156,0.0009856877],"domain_scores_gemma":[0.56445867,0.33977494,0.036023483,0.008951168,0.046549883,0.004241801],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.055215474,0.00044599638,0.0011073574,0.022192158,0.0010369602,0.007431784,0.0011859422,0.0008176776,0.0018592256],"category_scores_gemma":[0.32594454,0.0002020179,0.0011907994,0.045791566,0.0009895256,0.005407864,0.0026411968,0.0008473841,0.0006436178],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038295374,0.00032851842,0.66097337,0.0057292166,0.0034466772,0.00028875796,0.0025010195,0.0013614547,0.0019790959,0.004301279,0.011832163,0.30342898],"study_design_scores_gemma":[0.00020107198,0.0011930555,0.9577951,0.0019485609,0.0029148983,0.0004422105,0.0058254367,0.0037148634,0.003652519,0.0037760567,0.018345138,0.00019111455],"about_ca_topic_score_codex":0.009952013,"about_ca_topic_score_gemma":0.01759111,"teacher_disagreement_score":0.9778078,"about_ca_system_score_codex":0.0024598388,"about_ca_system_score_gemma":0.0036050244,"threshold_uncertainty_score":0.29201084},"labels":[],"label_agreement":null},{"id":"W4301502023","doi":"10.4018/978-1-60566-060-8.ch104","title":"Dimensions of UML Diagram Use","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland; University of Lethbridge","funders":"","keywords":"Applications of UML; Unified Modeling Language; UML tool; Class diagram; Computer science; Use Case Diagram; Software engineering; Communication diagram; Software; Programming language","score_opus":0.026261371365980495,"score_gpt":0.2634256375681515,"score_spread":0.23716426620217101,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4301502023","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13915017,0.01955114,0.16070372,0.013495331,0.00046625693,0.00052268675,0.0014967009,0.001181343,0.6634326],"genre_scores_gemma":[0.8096368,0.011366969,0.14564992,0.0013696211,0.00034976084,0.0007684271,0.0018817948,0.00056988985,0.028406864],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.974094,0.013812368,0.0017332555,0.0011663975,0.008681312,0.0005127582],"domain_scores_gemma":[0.96768713,0.021858945,0.0032021042,0.0025583715,0.003723108,0.0009702675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009247365,0.000543638,0.00024032165,0.006372172,0.0011623153,0.008127964,0.0007791701,0.0011863908,0.0043680067],"category_scores_gemma":[0.026233256,0.00035655248,0.00036236356,0.0071996246,0.0036531358,0.0082479,0.0041118036,0.001703038,0.0012551269],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005713789,0.00009393747,0.01667224,0.0007442339,0.000044073146,0.000260612,0.075802244,0.0008704586,0.0034752095,0.51566654,0.019206353,0.36710697],"study_design_scores_gemma":[0.000015119459,0.00006829197,0.01441048,0.0014674366,0.000024562098,0.0020656097,0.019751118,0.0017053271,0.0010847785,0.19345134,0.7658799,0.000075929805],"about_ca_topic_score_codex":0.0013958631,"about_ca_topic_score_gemma":0.0012516004,"teacher_disagreement_score":0.009247365,"about_ca_system_score_codex":0.0022925285,"about_ca_system_score_gemma":0.0016837969,"threshold_uncertainty_score":0.048905313},"labels":[],"label_agreement":null},{"id":"W4304889183","doi":"10.48550/arxiv.1807.04488","title":"Improved Query Reformulation for Concept Location using CodeRank and\\n Document Structures","year":2018,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Information retrieval; Query expansion; Source code; Sargable; Query optimization; Web query classification; Web search query; Baseline (sea); Task (project management); Software; Code (set theory); Term (time); Quality (philosophy); Data mining; Search engine; Programming language","score_opus":0.07155773527861449,"score_gpt":0.23551010109960796,"score_spread":0.16395236582099348,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4304889183","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08374373,0.0024553228,0.88666445,0.0012800945,0.0002625554,0.0008202142,0.0025426578,0.019045861,0.0031851905],"genre_scores_gemma":[0.2684757,0.0008076569,0.7128796,0.00038767638,0.0003267229,0.0003807806,0.010001186,0.0007832497,0.0059574475],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958657,0.0012691922,0.0004121049,0.0007078712,0.0014849969,0.00026009083],"domain_scores_gemma":[0.9914641,0.0039477483,0.0006169437,0.0016127262,0.0021201302,0.00023838413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023437666,0.0014887275,0.0019482232,0.0055896277,0.0010526488,0.0021066894,0.0020205537,0.0014238465,0.00517498],"category_scores_gemma":[0.013695029,0.00045483047,0.001297358,0.0046945075,0.0009452835,0.005795876,0.0021304574,0.002101174,0.0030436488],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000745049,0.00079104025,0.003002025,0.0009824855,0.000121272475,0.00037766772,0.001042952,0.036772504,0.061615266,0.018024277,0.044487555,0.8320379],"study_design_scores_gemma":[0.0002829668,0.00074515643,0.0022460374,0.00006289245,0.00014486711,0.0010027109,0.00097432174,0.8845253,0.060489316,0.017249929,0.032128464,0.00014803617],"about_ca_topic_score_codex":0.011812697,"about_ca_topic_score_gemma":0.013466141,"teacher_disagreement_score":0.011812697,"about_ca_system_score_codex":0.001562359,"about_ca_system_score_gemma":0.0033386536,"threshold_uncertainty_score":0.023487866},"labels":[],"label_agreement":null},{"id":"W4307177175","doi":"10.1145/3550355.3552421","title":"Machine learning-based incremental learning in interactive domain modelling","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Domain (mathematical analysis); Machine learning; Artificial intelligence; Subject-matter expert; Class (philosophy); Expert system","score_opus":0.016357742031070058,"score_gpt":0.25552701904909625,"score_spread":0.2391692770180262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4307177175","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07428171,0.00030779955,0.9157949,0.0004884178,0.00004493353,0.00034695712,0.00016124143,0.006527972,0.002046081],"genre_scores_gemma":[0.43643498,0.0001590239,0.5598249,0.00030449178,0.00003649024,0.00041577677,0.0007577243,0.00019146298,0.0018751567],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961964,0.001970249,0.00023495391,0.0007793991,0.0006303417,0.0001887267],"domain_scores_gemma":[0.97015846,0.023944203,0.0009463809,0.0024322057,0.002178114,0.00034068036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004630731,0.0014083231,0.001021454,0.001934436,0.0005796929,0.0013744896,0.0031917805,0.001513488,0.0017566343],"category_scores_gemma":[0.02423304,0.0006495106,0.0009911937,0.0012759896,0.00095296634,0.0030605318,0.0022991071,0.0023045135,0.0007791808],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048232166,0.0013786451,0.008558501,0.00050988473,0.00015097785,0.00036874405,0.0013414097,0.28135714,0.009844031,0.0059616202,0.0040522297,0.6859945],"study_design_scores_gemma":[0.000035770696,0.00012490785,0.0005540763,0.00002623502,0.00002445079,0.00007476977,0.0000857984,0.9849819,0.004037828,0.008178349,0.0018558063,0.000020019936],"about_ca_topic_score_codex":0.0046787793,"about_ca_topic_score_gemma":0.009197483,"teacher_disagreement_score":0.0046787793,"about_ca_system_score_codex":0.0013560047,"about_ca_system_score_gemma":0.0016660414,"threshold_uncertainty_score":0.02448994},"labels":[],"label_agreement":null},{"id":"W4307205996","doi":"10.1145/3550356.3561567","title":"An approach to build consistent software architecture diagrams using devops system descriptors","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Chicoutimi","funders":"","keywords":"Computer science; Reference architecture; Software architecture; Database-centric architecture; Consistency (knowledge bases); Architecture; Systems architecture; Software engineering; Software architecture description; Software system; DevOps; Use Case Diagram; Software; Class diagram; Artificial intelligence; Unified Modeling Language; Programming language","score_opus":0.04045819942232006,"score_gpt":0.27652862735586126,"score_spread":0.2360704279335412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4307205996","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027313635,0.00004642611,0.9937855,0.0001774182,0.000030319263,0.00035229587,0.00018367381,0.0014534092,0.0012395785],"genre_scores_gemma":[0.011327944,0.00006180731,0.98698425,0.00003093215,0.0000059636673,0.00024290216,0.00043752964,0.00021827768,0.0006904036],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99079424,0.0035820452,0.0014045375,0.0010527839,0.002925393,0.0002409921],"domain_scores_gemma":[0.97637707,0.007908461,0.0018160765,0.005477161,0.007870206,0.0005510396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01134361,0.0009982712,0.00065079576,0.0061713806,0.0014358822,0.00365479,0.0019429178,0.0013751993,0.0038704276],"category_scores_gemma":[0.032425325,0.0016146355,0.0016346687,0.0037831422,0.0012376792,0.0054755313,0.0039389455,0.0032079478,0.0015438932],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012315105,0.000353521,0.008411106,0.0013281074,0.00012482348,0.00061647984,0.007095199,0.0195908,0.024270365,0.23021016,0.013700466,0.69417584],"study_design_scores_gemma":[0.00022307037,0.0006251404,0.004216868,0.0010931472,0.00026189108,0.0017604012,0.0036137023,0.22134764,0.057713907,0.16201945,0.546778,0.00034673428],"about_ca_topic_score_codex":0.0034448009,"about_ca_topic_score_gemma":0.008248591,"teacher_disagreement_score":0.01134361,"about_ca_system_score_codex":0.0013470895,"about_ca_system_score_gemma":0.004824052,"threshold_uncertainty_score":0.05999148},"labels":[],"label_agreement":null},{"id":"W4307266360","doi":"10.3390/app122110750","title":"Game Development Topics: A Tag-Based Investigation on Game Development Stack Exchange","year":2022,"lang":"en","type":"article","venue":"Applied Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Perspective (graphical); Computer science; Game Developer; Video game development; Order (exchange); Video game; Face (sociological concept); Game design document; World Wide Web; Data science; Game design; Multimedia; Sociology; Artificial intelligence","score_opus":0.05939043689342571,"score_gpt":0.26174074229658445,"score_spread":0.20235030540315874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4307266360","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98311156,0.00034551535,0.010499019,0.00037807066,0.000050541123,0.00026339802,0.0006528343,0.00010353526,0.0045954934],"genre_scores_gemma":[0.9847986,0.00024012363,0.010058201,0.0002408757,0.00005721174,0.00038105441,0.0012579162,0.0001064492,0.0028595154],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99373615,0.0029019958,0.0005592119,0.00085032143,0.0013816098,0.0005706004],"domain_scores_gemma":[0.94887275,0.03722664,0.0051418277,0.0017364123,0.005373036,0.001649277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062973006,0.00038010097,0.0003864817,0.008191485,0.0017630106,0.0034269309,0.00062061154,0.0013417689,0.001387584],"category_scores_gemma":[0.03526471,0.00030064932,0.00035040535,0.005160894,0.0014217489,0.005267449,0.0037853369,0.0011754099,0.00058207085],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005539597,0.00029538388,0.28752106,0.0008297731,0.000043802323,0.001571447,0.59572023,0.00034272255,0.02098753,0.0072871824,0.0031743178,0.08167258],"study_design_scores_gemma":[0.00004958492,0.0005413405,0.33934605,0.00062218285,0.000101160185,0.0019137596,0.53973824,0.012125961,0.009956402,0.004683181,0.09070402,0.00021823059],"about_ca_topic_score_codex":0.0038463601,"about_ca_topic_score_gemma":0.006016215,"teacher_disagreement_score":0.008191485,"about_ca_system_score_codex":0.0015078959,"about_ca_system_score_gemma":0.0009459998,"threshold_uncertainty_score":0.033303738},"labels":[],"label_agreement":null},{"id":"W4307534208","doi":"10.1145/3550356.3561592","title":"Towards automatically extracting UML class diagrams from natural language specifications","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Class diagram; Computer science; Unified Modeling Language; UML tool; Applications of UML; Communication diagram; Programming language; Activity diagram; Object Constraint Language; Class (philosophy); Systems Modeling Language; Precision and recall; Software engineering; Natural language processing; Artificial intelligence; Software","score_opus":0.042661835709832466,"score_gpt":0.31214128784822237,"score_spread":0.2694794521383899,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4307534208","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15660948,0.004083794,0.7519103,0.0017379131,0.00030565457,0.0016321932,0.041828748,0.03700559,0.0048862286],"genre_scores_gemma":[0.10453639,0.00081444706,0.78660315,0.0003015044,0.000072888746,0.00075671443,0.1047952,0.0010454392,0.0010742758],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99128485,0.0036925722,0.000996228,0.0015957869,0.0022091737,0.0002214068],"domain_scores_gemma":[0.9522624,0.03250546,0.0034176689,0.0031792074,0.008049197,0.00058609573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071520093,0.0021278113,0.001059958,0.008769786,0.0008913947,0.0026324687,0.0016941403,0.0019511554,0.001587971],"category_scores_gemma":[0.035940994,0.00091451476,0.0016243932,0.0033728438,0.00060732325,0.0034868259,0.0019647577,0.0016471776,0.0024077066],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005178071,0.00083731976,0.030216374,0.0063632014,0.00041044265,0.0014113833,0.0035549372,0.01782257,0.08924437,0.014438198,0.094299756,0.74088365],"study_design_scores_gemma":[0.00067114586,0.0005337137,0.029636465,0.0011120398,0.00050870137,0.0024519628,0.0035186056,0.5058595,0.14786565,0.050063122,0.25747582,0.0003032969],"about_ca_topic_score_codex":0.0071830824,"about_ca_topic_score_gemma":0.010485855,"teacher_disagreement_score":0.008769786,"about_ca_system_score_codex":0.0013450306,"about_ca_system_score_gemma":0.0039923084,"threshold_uncertainty_score":0.037823856},"labels":[],"label_agreement":null},{"id":"W4307695669","doi":"10.48550/arxiv.2205.08116","title":"On the Use of Refactoring in Security Vulnerability Fixes: An Exploratory Study on Maven Libraries","year":2022,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Vulnerability (computing); Computer security; Security bug; Software engineering; Code (set theory); Software; Secure coding; Source code; Software security assurance; Programming language; Information security; Security service","score_opus":0.22230962567565407,"score_gpt":0.2409011261708669,"score_spread":0.018591500495212826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4307695669","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99624443,0.00015091622,0.0026530193,0.00007168236,0.0000019817119,0.00004502204,0.00006332767,0.000066750275,0.000702898],"genre_scores_gemma":[0.9899149,0.0002519792,0.008630547,0.00007113123,0.0000040424175,0.000070713235,0.00023766099,0.000113123206,0.0007058885],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9920791,0.0047107683,0.00051968574,0.0008565825,0.0013238053,0.0005100805],"domain_scores_gemma":[0.8739627,0.10425865,0.00926822,0.0053999107,0.006214006,0.0008965077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008538556,0.000503462,0.0003660118,0.0038541257,0.0010681123,0.0015220802,0.001111557,0.0009900498,0.00057111477],"category_scores_gemma":[0.054848727,0.0004794869,0.0004958173,0.0025957378,0.0015573989,0.0026776178,0.0018202758,0.0011213148,0.00023492923],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006691377,0.0019902533,0.565106,0.0009498273,0.00020218836,0.0046687066,0.16843933,0.0050886287,0.026432082,0.002868535,0.0012955852,0.2222897],"study_design_scores_gemma":[0.000091442525,0.003318845,0.7768966,0.0009035288,0.0004197286,0.0050215824,0.10326444,0.036949683,0.041803442,0.003726104,0.027299052,0.00030550396],"about_ca_topic_score_codex":0.0043387217,"about_ca_topic_score_gemma":0.009828227,"teacher_disagreement_score":0.008538556,"about_ca_system_score_codex":0.0011786069,"about_ca_system_score_gemma":0.0011734193,"threshold_uncertainty_score":0.045156777},"labels":[],"label_agreement":null},{"id":"W4307812665","doi":"10.1145/3569934","title":"What Is the Intended Usage Context of This Model? An Exploratory Study of Pre-Trained Models on Various Model Repositories","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Reuse; Benchmark (surveying); Leverage (statistics); Software engineering; Software; Machine learning; Artificial intelligence; Domain engineering; Code reuse; Context (archaeology); Software development; Software construction; Programming language","score_opus":0.1167313584876961,"score_gpt":0.33314403356947164,"score_spread":0.21641267508177553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4307812665","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6800328,0.0035100824,0.238929,0.006705214,0.000526867,0.0006059091,0.01618797,0.02316587,0.030336289],"genre_scores_gemma":[0.83940613,0.0011536148,0.1215656,0.0010314344,0.00010228091,0.0006759895,0.023027012,0.0073623057,0.005675585],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9879892,0.004557082,0.0010523193,0.0023148388,0.00328029,0.0008062051],"domain_scores_gemma":[0.943551,0.024017356,0.0018440903,0.022491828,0.007403007,0.00069276424],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012548388,0.0012431039,0.0011468759,0.0025113847,0.0013622813,0.0039128475,0.0030680366,0.0018609681,0.0035488245],"category_scores_gemma":[0.07779439,0.000896118,0.001693462,0.0030515115,0.0017557497,0.011906111,0.0023388434,0.0037600067,0.0021176946],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001788959,0.0013260284,0.18019573,0.0028059534,0.00078379654,0.0033674797,0.008119728,0.118346855,0.018500783,0.07009395,0.08310604,0.5115646],"study_design_scores_gemma":[0.00015814458,0.0006088787,0.029205216,0.00093645096,0.00037849226,0.0022134706,0.0033404615,0.80046177,0.03100394,0.03623262,0.09520467,0.00025589718],"about_ca_topic_score_codex":0.010522769,"about_ca_topic_score_gemma":0.014819662,"teacher_disagreement_score":0.9874516,"about_ca_system_score_codex":0.0020288725,"about_ca_system_score_gemma":0.003277014,"threshold_uncertainty_score":0.06636304},"labels":[],"label_agreement":null},{"id":"W4308623347","doi":"10.1145/3550356.3561566","title":"An investigation into the effect of cluster-based preprocessing on software migration","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Preprocessor; Partition (number theory); Obsolescence; Data mining; Artificial intelligence; Set (abstract data type); Software; Programming language; Transformation (genetics); Natural language processing; Theoretical computer science; Machine learning","score_opus":0.010831622836099985,"score_gpt":0.2678352767336775,"score_spread":0.2570036538975775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308623347","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76834345,0.0004577348,0.22021402,0.00070921687,0.00012975608,0.00057203916,0.00031810367,0.0042987247,0.0049569383],"genre_scores_gemma":[0.8362934,0.00018600606,0.16139887,0.00017672226,0.000025910911,0.000120784614,0.00037123862,0.0003499004,0.0010771813],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9972492,0.0010701854,0.00022991531,0.00050400826,0.00064527354,0.0003014411],"domain_scores_gemma":[0.96261716,0.021714788,0.002279561,0.009096799,0.003919903,0.00037175047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002827727,0.00085859007,0.00072931946,0.0011269656,0.0012021398,0.0013499574,0.0019333427,0.00078650785,0.0024686405],"category_scores_gemma":[0.021920318,0.000396627,0.0009495604,0.0022429586,0.0010189037,0.0017152872,0.0007522406,0.0013875689,0.0006118643],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039523416,0.0029004824,0.035225686,0.0010171972,0.00044472836,0.00077074417,0.0018132399,0.17793424,0.18020155,0.009306704,0.0033673667,0.58306587],"study_design_scores_gemma":[0.00028201027,0.0032282511,0.048255667,0.00009411506,0.000486263,0.0010176415,0.0014339158,0.5571207,0.36647436,0.010868521,0.010604554,0.00013400977],"about_ca_topic_score_codex":0.0068213735,"about_ca_topic_score_gemma":0.006942072,"teacher_disagreement_score":0.0068213735,"about_ca_system_score_codex":0.0012037387,"about_ca_system_score_gemma":0.0016402409,"threshold_uncertainty_score":0.014954686},"labels":[],"label_agreement":null},{"id":"W4308627662","doi":"10.1145/3549035.3561182","title":"An exploratory study on the relationship of smells and design issues with software vulnerabilities","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code smell; Secure coding; Computer science; Vulnerability (computing); Software; Computer security; Exploratory research; Confidentiality; Software design; Software security assurance; Security bug; Software engineering; Code (set theory); Software development; Software quality; Information security","score_opus":0.08005182148217244,"score_gpt":0.299831598857228,"score_spread":0.21977977737505555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308627662","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99880147,0.00004851293,0.00044055938,0.00005977861,0.0000013637614,0.000036725418,0.000056237408,0.0000069684584,0.000548568],"genre_scores_gemma":[0.9984101,0.00007497803,0.0010218609,0.000032470736,0.0000039630745,0.000048617516,0.000084175066,0.000005746526,0.0003179756],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9959455,0.0016831737,0.00038031308,0.00038172735,0.0012388312,0.0003705425],"domain_scores_gemma":[0.8796227,0.091366135,0.017814964,0.002522006,0.0065504913,0.002123757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060831816,0.00040446798,0.00034671277,0.0026784204,0.00080149184,0.0013553604,0.0005804075,0.00076213083,0.0011196556],"category_scores_gemma":[0.04556682,0.0004241871,0.0004797644,0.0019400186,0.0011347028,0.0024399213,0.0016598002,0.0013335731,0.00019909958],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025167246,0.0007747549,0.90981174,0.00026300622,0.00007243606,0.0023978013,0.056818534,0.0003555074,0.006020979,0.00054669735,0.0003413144,0.022345463],"study_design_scores_gemma":[0.000011166734,0.0007745087,0.9552295,0.00007607042,0.0000393984,0.0012537724,0.0366022,0.0016050688,0.0018549658,0.0004753975,0.0020336416,0.000044276538],"about_ca_topic_score_codex":0.0011182895,"about_ca_topic_score_gemma":0.00206798,"teacher_disagreement_score":0.0060831816,"about_ca_system_score_codex":0.00097146543,"about_ca_system_score_gemma":0.000843491,"threshold_uncertainty_score":0.03217131},"labels":[],"label_agreement":null},{"id":"W4308632594","doi":"10.1145/3550356.3552372","title":"Automated, traceable, and interactive domain modelling","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Traceability; Computer science; Domain (mathematical analysis); Warrant; Domain model; Domain analysis; Software engineering; Data science; Domain knowledge; Programming language; Software; Software development","score_opus":0.01629811112062965,"score_gpt":0.2578718476089686,"score_spread":0.24157373648833894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308632594","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022730264,0.000037650952,0.98741573,0.00029408638,0.000020340533,0.00014492629,0.00016976871,0.007593968,0.0020505718],"genre_scores_gemma":[0.06634505,0.00018603021,0.92445475,0.000259009,0.00002272142,0.00028587226,0.0012973427,0.002179757,0.0049694],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98987997,0.0037816528,0.00076669635,0.0017555897,0.0033015036,0.000514658],"domain_scores_gemma":[0.9646091,0.016535738,0.0018557778,0.012980752,0.0035003521,0.0005182856],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008195391,0.0021005494,0.0010875957,0.0032159293,0.0016984936,0.004208348,0.00513719,0.0032398214,0.010905892],"category_scores_gemma":[0.035126377,0.0017520448,0.0022420431,0.0027114116,0.0023079463,0.008022719,0.009357875,0.004976477,0.0044753193],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047633538,0.001113519,0.0056181536,0.0011335239,0.0002437744,0.0014848569,0.006609413,0.14270683,0.03568183,0.17304908,0.029355438,0.6025273],"study_design_scores_gemma":[0.000086787404,0.00009875109,0.0005988017,0.00024968793,0.000066534216,0.0006587522,0.0007280098,0.77646613,0.015761,0.11777328,0.08741279,0.0000994679],"about_ca_topic_score_codex":0.009252535,"about_ca_topic_score_gemma":0.02183179,"teacher_disagreement_score":0.010905892,"about_ca_system_score_codex":0.0022036268,"about_ca_system_score_gemma":0.00556279,"threshold_uncertainty_score":0.043341875},"labels":[],"label_agreement":null},{"id":"W4308643006","doi":"10.1145/3540250.3549079","title":"Accurate method and variable tracking in commit history","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 30th ACM Joint European Software Engineering Conference and Symposium on the Foundations of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Commit; Computer science; Oracle; Software evolution; Variable (mathematics); Program comprehension; Tracking (education); Code refactoring; Precision and recall; Software engineering; Software; Programming language; Artificial intelligence; Software development; Software system; Database; Software construction","score_opus":0.03538692714982885,"score_gpt":0.24103391154454729,"score_spread":0.20564698439471843,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308643006","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17550221,0.004571551,0.6828382,0.00096927746,0.0006093038,0.0005067315,0.0076487404,0.11732124,0.01003273],"genre_scores_gemma":[0.6093743,0.0008672401,0.36053264,0.00045532524,0.0001312837,0.0003559633,0.014795085,0.008318957,0.005169234],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9835414,0.002999511,0.001925898,0.0033273723,0.007527105,0.00067864574],"domain_scores_gemma":[0.93159837,0.027830891,0.006660401,0.019567776,0.013378041,0.00096457335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008541136,0.001115477,0.0010030078,0.004493838,0.0008855804,0.0038909784,0.0023425526,0.0015557322,0.0025760015],"category_scores_gemma":[0.09432251,0.0011669613,0.00079347315,0.002570082,0.0009647153,0.006898315,0.0031868627,0.0021838488,0.002017198],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010600864,0.00025967622,0.12446445,0.0017741781,0.00020947026,0.00064202084,0.002610827,0.030528983,0.01889826,0.015427609,0.038643185,0.7654813],"study_design_scores_gemma":[0.000300175,0.0008405627,0.06368699,0.0012512624,0.0004017358,0.0021130524,0.0011689672,0.6002064,0.12448993,0.047846254,0.1572655,0.000429307],"about_ca_topic_score_codex":0.007966516,"about_ca_topic_score_gemma":0.009087665,"teacher_disagreement_score":0.008541136,"about_ca_system_score_codex":0.0008392176,"about_ca_system_score_gemma":0.0031263623,"threshold_uncertainty_score":0.045170367},"labels":[],"label_agreement":null},{"id":"W4308643044","doi":"10.1145/3540250.3549112","title":"PaReco: patched clones and missed patches among the divergent variants of a software family","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 30th ACM Joint European Software Engineering Conference and Symposium on the Foundations of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Wetenschappelijk Onderzoek; Vlaamse regering; Fonds De La Recherche Scientifique - FNRS","keywords":"Patched; Software; Computer science; Biology; Genetics; Gene; Programming language; Hedgehog signaling pathway","score_opus":0.029019680987102957,"score_gpt":0.21124654982490057,"score_spread":0.18222686883779762,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308643044","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8817667,0.0027256017,0.086270384,0.00040422744,0.00018253124,0.00031868473,0.010179217,0.015252271,0.0029003862],"genre_scores_gemma":[0.8723934,0.0006230253,0.08613313,0.00021343377,0.00012984558,0.00026700477,0.034862246,0.0013790543,0.0039988146],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9940069,0.0006878382,0.00043854682,0.002876668,0.0015994579,0.00039062955],"domain_scores_gemma":[0.97270614,0.012855979,0.0049930993,0.005891905,0.0026970787,0.00085575273],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034981123,0.0012411589,0.0010500251,0.006416282,0.0009235809,0.0015685689,0.0020084374,0.00168831,0.0010693341],"category_scores_gemma":[0.02416082,0.0006372839,0.0010285944,0.003483835,0.0008957978,0.0029453293,0.0023645184,0.00090756395,0.0008229194],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010203316,0.00041767082,0.62970775,0.0015656131,0.00057187193,0.0048837117,0.0026955607,0.007907302,0.028824607,0.0022993288,0.0387083,0.2813979],"study_design_scores_gemma":[0.00022632998,0.0009603026,0.6042883,0.00045554267,0.00065666216,0.021515572,0.002464058,0.2565909,0.041442093,0.007902266,0.06324815,0.0002497928],"about_ca_topic_score_codex":0.0038391503,"about_ca_topic_score_gemma":0.007698593,"teacher_disagreement_score":0.006416282,"about_ca_system_score_codex":0.00048066227,"about_ca_system_score_gemma":0.00082799786,"threshold_uncertainty_score":0.01849997},"labels":[],"label_agreement":null},{"id":"W4308643083","doi":"10.1145/3540250.3558952","title":"Leveraging test plan quality to improve code review efficacy","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 30th ACM Joint European Software Engineering Conference and Symposium on the Foundations of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Code review; Leverage (statistics); Documentation; Code coverage; Test case; Source code; Natural language; Transformer; Test (biology); Software engineering; Software quality; Artificial intelligence; Data science; Information retrieval; Natural language processing; Machine learning; Programming language; Software; Software development; Engineering","score_opus":0.04372462117162971,"score_gpt":0.25991377766099844,"score_spread":0.21618915648936873,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308643083","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.72164613,0.0028950037,0.26090756,0.0020316048,0.00019516988,0.0010507505,0.0011718443,0.0051806336,0.0049211513],"genre_scores_gemma":[0.97214365,0.00019503123,0.02560011,0.0001746827,0.000053491684,0.00012345547,0.0009709038,0.00013418202,0.0006044214],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97267735,0.013475835,0.0026688974,0.0041827583,0.0063660583,0.00062920444],"domain_scores_gemma":[0.71560705,0.227215,0.02209639,0.008576124,0.024492119,0.0020132991],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.037308883,0.0010605196,0.0011416646,0.005231675,0.0005251719,0.0029533599,0.001315919,0.0012410117,0.0011264601],"category_scores_gemma":[0.19253382,0.00043802432,0.00093266665,0.0020579747,0.0008385202,0.004814939,0.0014484188,0.0015210101,0.0007522406],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017577197,0.0011971709,0.39508983,0.0016537542,0.0010055504,0.0003291106,0.0037724136,0.06458274,0.015514877,0.0021407434,0.008379116,0.5045771],"study_design_scores_gemma":[0.00015051024,0.0014184299,0.10709804,0.00015593364,0.00045299003,0.00032069677,0.00065195706,0.864974,0.014522553,0.0061267433,0.003984049,0.00014410244],"about_ca_topic_score_codex":0.004338627,"about_ca_topic_score_gemma":0.006764498,"teacher_disagreement_score":0.037308883,"about_ca_system_score_codex":0.0018379678,"about_ca_system_score_gemma":0.0024430566,"threshold_uncertainty_score":0.19731063},"labels":[],"label_agreement":null},{"id":"W4308648311","doi":"10.1145/3540250.3549156","title":"Automated unearthing of dangerous issue reports","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 30th ACM Joint European Software Engineering Conference and Symposium on the Foundations of Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Huawei Technologies (Canada)","funders":"Fundamental Research Funds for the Central Universities","keywords":"Vulnerability (computing); Operationalization; Computer science; Computer security; Vulnerability assessment; Process (computing); Best practice; Internet privacy; Political science; Psychology; Psychological resilience","score_opus":0.019585129241062014,"score_gpt":0.22313671679793548,"score_spread":0.20355158755687347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308648311","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.59018123,0.01372614,0.21501268,0.003016746,0.0009139775,0.0006893803,0.05787754,0.10114263,0.017439617],"genre_scores_gemma":[0.7050987,0.0024317156,0.17535765,0.0004736619,0.0003076415,0.0002418679,0.10661861,0.0013914237,0.008078674],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956957,0.00074000546,0.00048211752,0.0012252826,0.001583437,0.00027339318],"domain_scores_gemma":[0.9830487,0.006350365,0.004209608,0.0030005313,0.0029262954,0.00046443407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036841808,0.0018498931,0.0011226695,0.0111095905,0.00046282256,0.0021798715,0.0020766882,0.0013411102,0.0016559646],"category_scores_gemma":[0.021267353,0.0006047296,0.0010932315,0.003987365,0.00046706572,0.0039105713,0.002810986,0.0013353786,0.0024160957],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042840274,0.000284522,0.111018434,0.0016583768,0.00026589882,0.0015523394,0.0012629086,0.015667787,0.010811753,0.0029425356,0.09032628,0.76378083],"study_design_scores_gemma":[0.00013747995,0.0004129454,0.09320337,0.0010981014,0.0005925146,0.0034709298,0.0023767196,0.6110734,0.0729708,0.014855493,0.19957697,0.00023133915],"about_ca_topic_score_codex":0.0044030403,"about_ca_topic_score_gemma":0.0059230425,"teacher_disagreement_score":0.0111095905,"about_ca_system_score_codex":0.00059025455,"about_ca_system_score_gemma":0.0018802823,"threshold_uncertainty_score":0.019484043},"labels":[],"label_agreement":null},{"id":"W4309382233","doi":"10.1007/s10664-022-10244-0","title":"Developer discussion topics on the adoption and barriers of low code software development platforms","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Agile software development; Computer science; Software development; Personalization; Software; Software development process; Software engineering; World Wide Web; Data science; Engineering","score_opus":0.02297163341640662,"score_gpt":0.2529683407359211,"score_spread":0.2299967073195145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4309382233","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.74969435,0.0043994184,0.024970649,0.1063324,0.0046021775,0.0014973372,0.0029465847,0.00072304026,0.10483411],"genre_scores_gemma":[0.92885864,0.003002153,0.008489207,0.008487807,0.0019000736,0.0018685807,0.0014911807,0.00041137714,0.045491006],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.99484575,0.0024040297,0.0003517477,0.00043889767,0.0010369663,0.00092255557],"domain_scores_gemma":[0.91411716,0.059727855,0.0056775706,0.0016836877,0.0114447465,0.007348955],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011925642,0.0005563123,0.00041328973,0.0030358315,0.0039628083,0.0024557938,0.0008980138,0.003031293,0.018699478],"category_scores_gemma":[0.06658877,0.00041048016,0.0006039894,0.0019513088,0.00088915817,0.0046640886,0.0042990716,0.003120569,0.0016938668],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019464245,0.0019115545,0.20854053,0.0042004897,0.00014303274,0.00256719,0.22110441,0.0013995738,0.01824305,0.036366016,0.14695407,0.35662362],"study_design_scores_gemma":[0.00020965192,0.0014403147,0.2281364,0.0033742655,0.00017661584,0.0008198894,0.17245144,0.0022525585,0.010087751,0.01528695,0.5655598,0.00020435051],"about_ca_topic_score_codex":0.0023786228,"about_ca_topic_score_gemma":0.0043911813,"teacher_disagreement_score":0.98807436,"about_ca_system_score_codex":0.003054953,"about_ca_system_score_gemma":0.004682285,"threshold_uncertainty_score":0.06306958},"labels":[],"label_agreement":null},{"id":"W4309529433","doi":"10.1145/3571848","title":"On the Discoverability of npm Vulnerabilities in Node.js Projects","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Concordia University","funders":"","keywords":"Discoverability; Computer science; Dependency (UML); Vulnerability (computing); Computer security; World Wide Web; Software engineering","score_opus":0.0900358588716811,"score_gpt":0.3147622559724983,"score_spread":0.2247263971008172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4309529433","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99775994,0.000086263964,0.0007491445,0.00012373536,0.0000016410868,0.000017942102,0.00019260202,0.000026732123,0.0010419834],"genre_scores_gemma":[0.9982805,0.0000978716,0.0010451672,0.000018200748,0.0000045027246,0.000021995662,0.00030867133,0.000015952433,0.00020716396],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9916524,0.0019843203,0.00089455943,0.0012555424,0.003366329,0.0008467596],"domain_scores_gemma":[0.6846758,0.23171626,0.061443567,0.008010489,0.011465554,0.0026884477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011836533,0.0003482178,0.00024109783,0.008141704,0.0011247993,0.00151825,0.00078213523,0.00088373653,0.0015822961],"category_scores_gemma":[0.09871123,0.00035961354,0.0004966157,0.0048430124,0.0020736223,0.004311473,0.0024530524,0.0013171874,0.00025859452],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000076131626,0.00011746355,0.98090404,0.00008018626,0.000036891604,0.0003514042,0.0034550594,0.0015135262,0.0008238605,0.0007058804,0.00029042966,0.011645206],"study_design_scores_gemma":[0.0000048006145,0.00012097512,0.98674345,0.00006219194,0.00003124393,0.0004988566,0.0030316105,0.0069321427,0.0009025509,0.0008802328,0.0007698279,0.000022007876],"about_ca_topic_score_codex":0.009378598,"about_ca_topic_score_gemma":0.017957311,"teacher_disagreement_score":0.011836533,"about_ca_system_score_codex":0.0012927861,"about_ca_system_score_gemma":0.0014421304,"threshold_uncertainty_score":0.06259829},"labels":[],"label_agreement":null},{"id":"W4309896433","doi":"10.1007/s10664-022-10242-2","title":"Automatic prediction of rejected edits in Stack Overflow","year":2022,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Global Institute for Water Security, University of Saskatchewan","keywords":"Computer science; Quality (philosophy); World Wide Web; Information retrieval; Software; Precision and recall; Artificial intelligence; Programming language","score_opus":0.0211596723309598,"score_gpt":0.2590476734414302,"score_spread":0.2378880011104704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4309896433","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9911809,0.00023026661,0.0048055043,0.00007143606,0.00006948076,0.0000288652,0.0016516056,0.001240678,0.0007212933],"genre_scores_gemma":[0.9900726,0.00006509769,0.0056272335,0.000028760753,0.000055634984,0.000013603354,0.0032540346,0.000114926,0.00076816085],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997964,0.00032349234,0.0002081606,0.0004908502,0.0008201687,0.00019336019],"domain_scores_gemma":[0.9576629,0.023792336,0.0074135168,0.0023055088,0.0069742906,0.0018515304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016438977,0.00062266213,0.00049685297,0.0046891384,0.0006253858,0.0012388185,0.0010590224,0.0012077418,0.0017331204],"category_scores_gemma":[0.024771202,0.00025347542,0.00043670516,0.0015328827,0.00028594973,0.0011816488,0.00068522914,0.0010602408,0.0009816389],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021389222,0.0007142964,0.8409462,0.00035329864,0.00026722456,0.0013179334,0.0005153657,0.01021805,0.017800303,0.00072409335,0.007562227,0.117442146],"study_design_scores_gemma":[0.00010535357,0.0006823094,0.54398066,0.00010520737,0.00032823626,0.0017977469,0.0005054016,0.42224658,0.024708005,0.00156569,0.0038503634,0.00012441736],"about_ca_topic_score_codex":0.004825873,"about_ca_topic_score_gemma":0.007856971,"teacher_disagreement_score":0.004825873,"about_ca_system_score_codex":0.00035891056,"about_ca_system_score_gemma":0.0009512057,"threshold_uncertainty_score":0.009595513},"labels":[],"label_agreement":null},{"id":"W4311087681","doi":"10.1016/j.jocs.2022.101928","title":"A broad approach to expert detection using syntactic and semantic social networks analysis in the context of Global Software Development","year":2022,"lang":"en","type":"article","venue":"Journal of Computational Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Fundação de Amparo à Pesquisa do Estado de Minas Gerais; Conselho Nacional de Desenvolvimento Científico e Tecnológico","keywords":"Computer science; Ontology; Task (project management); Social network analysis; Data science; Context (archaeology); Social network (sociolinguistics); Software development; Software; Knowledge management; Software engineering; Artificial intelligence; World Wide Web; Social media","score_opus":0.024940242092239495,"score_gpt":0.30004047692677327,"score_spread":0.2751002348345338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4311087681","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10562545,0.0009216478,0.87953544,0.0021259214,0.00005934186,0.00031688085,0.00065127097,0.00051431503,0.010249673],"genre_scores_gemma":[0.7077038,0.00068568107,0.287057,0.00032815922,0.00014969177,0.00022315283,0.0007628912,0.00009316794,0.0029964573],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953284,0.0019464782,0.00025401937,0.0010000129,0.0011991652,0.0002718123],"domain_scores_gemma":[0.98880076,0.0072974768,0.00084134575,0.0010200236,0.0016597467,0.00038050232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044365404,0.0008624298,0.0008694532,0.0091048265,0.0020757972,0.0038658727,0.0016295353,0.0024147606,0.0019159337],"category_scores_gemma":[0.017385721,0.00053525274,0.0009788302,0.0043974044,0.002179883,0.00686808,0.0036113486,0.0019562272,0.00069761975],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044345311,0.0010776296,0.091245174,0.0008560566,0.0006599176,0.0020587342,0.008195766,0.04997971,0.03189123,0.27071074,0.0075215143,0.5353601],"study_design_scores_gemma":[0.000029295476,0.00014860395,0.024719676,0.0002151483,0.000228435,0.0011466979,0.0052907662,0.56933284,0.009350004,0.37443233,0.014991115,0.000115134455],"about_ca_topic_score_codex":0.0039749364,"about_ca_topic_score_gemma":0.006417775,"teacher_disagreement_score":0.0091048265,"about_ca_system_score_codex":0.0009863174,"about_ca_system_score_gemma":0.0023153035,"threshold_uncertainty_score":0.023462951},"labels":[],"label_agreement":null},{"id":"W4312055758","doi":"10.1016/j.jss.2022.111590","title":"A study of update request comments in Stack Overflow answer posts","year":2022,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; JavaScript; Python (programming language); Classifier (UML); Questions and answers; Set (abstract data type); Information retrieval; Point (geometry); Machine learning; Ask price; World Wide Web; Artificial intelligence; Data mining; Programming language","score_opus":0.02101511423514871,"score_gpt":0.2736684938976722,"score_spread":0.2526533796625235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312055758","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9955023,0.00007597673,0.00048271118,0.0001571126,0.000027911945,0.00008244302,0.00046917106,0.00008280469,0.0031196135],"genre_scores_gemma":[0.994269,0.00008606583,0.0008614179,0.00008966875,0.000052299933,0.00006590343,0.0010394792,0.00006799414,0.0034682332],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9950257,0.0017169723,0.00035848035,0.00034568398,0.0021175176,0.00043563131],"domain_scores_gemma":[0.7950195,0.15570381,0.021631682,0.0032140287,0.021388138,0.003042717],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0024052248,0.00031916102,0.00030571216,0.0036850998,0.0019085475,0.0023876284,0.0006927107,0.0014647215,0.005216071],"category_scores_gemma":[0.069735885,0.0003181618,0.0002626845,0.0034403217,0.00069922354,0.0025167705,0.00078212365,0.0012864678,0.00185254],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029996682,0.0027692993,0.84506,0.00072273216,0.00012173179,0.003062385,0.040828533,0.0009557175,0.015102322,0.0023235546,0.007852014,0.07820212],"study_design_scores_gemma":[0.000080692575,0.0017234433,0.89969236,0.00026837058,0.00020047382,0.0016363477,0.05422427,0.013836726,0.009530574,0.000753847,0.01791563,0.00013728051],"about_ca_topic_score_codex":0.011742432,"about_ca_topic_score_gemma":0.013955127,"teacher_disagreement_score":0.99761236,"about_ca_system_score_codex":0.0014630745,"about_ca_system_score_gemma":0.0016678132,"threshold_uncertainty_score":0.023348212},"labels":[],"label_agreement":null},{"id":"W4312234739","doi":"10.1109/icsme55016.2022.00053","title":"Integrating Software Issue Tracking and Traceability Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Traceability; Requirements traceability; Computer science; Process (computing); Software engineering; Systems engineering; Process management; Software; Risk analysis (engineering); Software development; Engineering; Business; Requirement","score_opus":0.02790534711123356,"score_gpt":0.26921900406896726,"score_spread":0.2413136569577337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312234739","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004769979,0.00010032977,0.986193,0.0008183068,0.00008691193,0.00029181052,0.00006268803,0.0040987846,0.0035782147],"genre_scores_gemma":[0.12997575,0.00038945972,0.8653222,0.00038389076,0.000106977706,0.00047421793,0.00040462194,0.00046966696,0.0024733066],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9792188,0.0064759385,0.0022307856,0.0022080413,0.009054883,0.0008116114],"domain_scores_gemma":[0.9398634,0.029103031,0.0046660313,0.017673546,0.0077406205,0.0009533423],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018391898,0.0014872788,0.0010659101,0.004909138,0.0015650562,0.008188808,0.004624796,0.0032458818,0.005061367],"category_scores_gemma":[0.06234352,0.001715636,0.0017089979,0.0026724301,0.002966317,0.021008799,0.0069010044,0.0061190766,0.0013141361],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002888318,0.0011830357,0.0049283877,0.0006862588,0.00020260877,0.0006222121,0.0016595391,0.07435103,0.010778255,0.47842526,0.006403496,0.42047113],"study_design_scores_gemma":[0.00016524293,0.0004575796,0.00091294944,0.0004976669,0.00016215692,0.0006888511,0.00037899267,0.6107009,0.023331609,0.30314165,0.059361465,0.00020091709],"about_ca_topic_score_codex":0.006363694,"about_ca_topic_score_gemma":0.003922514,"teacher_disagreement_score":0.018391898,"about_ca_system_score_codex":0.0028481146,"about_ca_system_score_gemma":0.0066450653,"threshold_uncertainty_score":0.09726679},"labels":[],"label_agreement":null},{"id":"W4312244873","doi":"10.1145/3524842.3527975","title":"An empirical study on maintainable method size in Java","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of British Columbia","funders":"","keywords":"Computer science; Java; Software maintenance; Metric (unit); Source lines of code; Empirical research; Software metric; Maintainability; Software engineering; Readability; Software; Software development; Programming language; Software quality; Statistics","score_opus":0.03505638865765372,"score_gpt":0.38877599024061943,"score_spread":0.3537196015829657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312244873","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9986902,0.000084153966,0.0002715886,0.000105108134,0.0000029781245,0.000015570246,0.00013376572,0.000011655182,0.00068504346],"genre_scores_gemma":[0.99867886,0.000060254883,0.0005490682,0.00004320303,0.0000063851426,0.000032806038,0.00031090016,0.000016284039,0.0003022656],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.993346,0.0027521048,0.0006660958,0.0008822169,0.0020442994,0.000309345],"domain_scores_gemma":[0.618799,0.2865331,0.05269718,0.010844496,0.02565464,0.0054716277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009577751,0.00024149098,0.00018977188,0.001716434,0.00054871873,0.00131117,0.001033396,0.0006794941,0.0012788117],"category_scores_gemma":[0.13170545,0.00029696358,0.00026346731,0.0019974625,0.0011806694,0.0028950886,0.00067911536,0.0015084888,0.00032386166],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020382325,0.00074042263,0.9715314,0.00014368737,0.00006116931,0.00024016108,0.0073413816,0.000479865,0.00095351454,0.00033345452,0.00077642116,0.017194703],"study_design_scores_gemma":[0.000016621485,0.0004094582,0.9902789,0.000049755938,0.00002496792,0.00030487988,0.0038925805,0.0024867805,0.0005656093,0.00021296409,0.0017397304,0.000017748524],"about_ca_topic_score_codex":0.0038240694,"about_ca_topic_score_gemma":0.006185039,"teacher_disagreement_score":0.009577751,"about_ca_system_score_codex":0.0009002069,"about_ca_system_score_gemma":0.00079619786,"threshold_uncertainty_score":0.050652623},"labels":[],"label_agreement":null},{"id":"W4312307983","doi":"10.1145/3524842.3527996","title":"ApacheJIT","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Software; Software bug; Training set; Artificial intelligence; Data modeling; Software engineering; Machine learning; Data mining; Programming language","score_opus":0.01820738705578956,"score_gpt":0.24988959650723366,"score_spread":0.2316822094514441,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312307983","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025665231,0.00095999654,0.004687803,0.0004633948,0.00041308784,0.0002735246,0.92526686,0.031734936,0.010535298],"genre_scores_gemma":[0.015906164,0.0001694296,0.004498757,0.00014839204,0.000039996958,0.00017410974,0.97619617,0.00095699285,0.001909884],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979146,0.00022282125,0.00019775805,0.0005520749,0.0007815778,0.00033116253],"domain_scores_gemma":[0.9969489,0.00039983107,0.0003374603,0.0010933921,0.0008083949,0.00041211257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013274306,0.0018315075,0.0009854209,0.0032090156,0.00088128913,0.0017660267,0.0024670756,0.0013929562,0.01017662],"category_scores_gemma":[0.005201681,0.0007075604,0.0012483858,0.0045711687,0.00054865464,0.0024642304,0.0023150833,0.001993887,0.01852128],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081203616,0.0002986818,0.019881943,0.00093980785,0.00018454785,0.00039185205,0.0001749842,0.0038434272,0.0025506874,0.00227777,0.93332046,0.035323773],"study_design_scores_gemma":[0.0005186724,0.0003809426,0.061972305,0.00027324687,0.00014492436,0.0008785501,0.00026968692,0.023893679,0.006397744,0.006010956,0.8990496,0.00020981095],"about_ca_topic_score_codex":0.016218606,"about_ca_topic_score_gemma":0.027154345,"teacher_disagreement_score":0.016218606,"about_ca_system_score_codex":0.00066491327,"about_ca_system_score_gemma":0.0017760138,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4312345958","doi":"10.1145/3510455.3512793","title":"Supporting program comprehension by generating abstract code summary tree","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Codebase; Program comprehension; Tree (set theory); Source code; Node (physics); Software; Code (set theory); Software maintenance; Cluster analysis; Programming language; Theoretical computer science; Software system; Artificial intelligence","score_opus":0.025275417511071867,"score_gpt":0.31248482074930567,"score_spread":0.2872094032382338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312345958","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04582767,0.00025819518,0.88671374,0.00039019177,0.000087814245,0.0006723615,0.005545164,0.05760964,0.002895228],"genre_scores_gemma":[0.11662437,0.00025765845,0.8573685,0.00016953098,0.000050713603,0.00048677495,0.017811276,0.0042486265,0.0029825633],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99871564,0.00026314077,0.00013896674,0.00030418357,0.00051404536,0.000064005006],"domain_scores_gemma":[0.9901329,0.0046742046,0.0009783342,0.001266486,0.0027257958,0.00022232249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012783002,0.0013657289,0.00069177087,0.0034639174,0.00056931336,0.0013782405,0.0012176284,0.00095811195,0.0048915152],"category_scores_gemma":[0.012570712,0.00055240595,0.00088044396,0.0023188882,0.0003796228,0.0028755248,0.0011604798,0.0010592338,0.0023927882],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006560263,0.0004373634,0.0067283143,0.0018383716,0.00008433982,0.0011861499,0.005219264,0.021746473,0.09166191,0.009339362,0.05133351,0.80976903],"study_design_scores_gemma":[0.00025611033,0.0007836722,0.009373758,0.00033845345,0.0002876559,0.0013313079,0.002001305,0.6376314,0.13568407,0.0349939,0.17708576,0.00023264026],"about_ca_topic_score_codex":0.0026761717,"about_ca_topic_score_gemma":0.004251017,"teacher_disagreement_score":0.0048915152,"about_ca_system_score_codex":0.00052680785,"about_ca_system_score_gemma":0.0014455534,"threshold_uncertainty_score":0.01636374},"labels":[],"label_agreement":null},{"id":"W4312422135","doi":"10.1145/3510457.3513039","title":"How does code reviewing feedback evolve?","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Polytechnique Montréal; McGill University","funders":"Mitacs","keywords":"Code review; Computer science; Code (set theory); Context (archaeology); Process (computing); Generalizability theory; Premise; Replicate; Static program analysis; Software; Knowledge management; Software development; Software engineering; Psychology","score_opus":0.027171326006801826,"score_gpt":0.2577381481836016,"score_spread":0.23056682217679975,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312422135","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9229411,0.0054344516,0.022702765,0.019841807,0.0005995948,0.0005398904,0.0010236765,0.0012441304,0.025672523],"genre_scores_gemma":[0.98924214,0.0009271995,0.0058167735,0.0011096479,0.00024222878,0.0001500537,0.00039590098,0.00026350364,0.0018526013],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.8887028,0.053862937,0.008357759,0.013423817,0.031668454,0.003984251],"domain_scores_gemma":[0.34792754,0.37895963,0.10667792,0.03378427,0.1236395,0.009011141],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07908302,0.0006125783,0.0011414167,0.009485941,0.0027414437,0.010195134,0.0023798007,0.0031440703,0.0020432733],"category_scores_gemma":[0.48650917,0.0010706414,0.0008583835,0.0073964684,0.002553752,0.012604352,0.003191952,0.0018941667,0.001901614],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031167612,0.00034394374,0.598469,0.0008744408,0.00030811387,0.00036738714,0.035966855,0.0017105443,0.0021602921,0.0032308123,0.009837877,0.34641907],"study_design_scores_gemma":[0.000065554836,0.0006767393,0.90796375,0.0013329334,0.00023630243,0.0011327005,0.025331113,0.012267567,0.0027063687,0.009036086,0.038940623,0.00031036022],"about_ca_topic_score_codex":0.012088404,"about_ca_topic_score_gemma":0.014348888,"teacher_disagreement_score":0.920917,"about_ca_system_score_codex":0.00529045,"about_ca_system_score_gemma":0.0074632145,"threshold_uncertainty_score":0.41823596},"labels":[],"label_agreement":null},{"id":"W4312437398","doi":"10.1145/3524610.3527872","title":"Deep API learning revisited","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Python (programming language); Java; Documentation; Encoder; Source code; Architecture; Artificial intelligence; Deep learning; Preprocessor; Task (project management); Programming language; Machine learning; Information retrieval; Operating system","score_opus":0.013438660287173549,"score_gpt":0.2508100375218349,"score_spread":0.23737137723466134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312437398","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.088629276,0.005377094,0.8568906,0.006578189,0.00073748385,0.0001347732,0.0021877636,0.016359938,0.023104878],"genre_scores_gemma":[0.7826079,0.0021383066,0.1759753,0.0022117496,0.00027937413,0.0001774712,0.0062720617,0.0010751871,0.029262673],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991204,0.00020699468,0.00004905782,0.0002896114,0.00019787202,0.00013605501],"domain_scores_gemma":[0.99823976,0.00050891086,0.000085314765,0.0006148773,0.0004493188,0.00010177261],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013475477,0.0011113659,0.00081672287,0.00082885,0.00042524503,0.0016445778,0.002380376,0.0012371499,0.005683317],"category_scores_gemma":[0.006264578,0.00045802273,0.00081025297,0.0012186685,0.00070904644,0.0035254764,0.0016149299,0.0033932098,0.0030161277],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003452201,0.00029279708,0.0047998345,0.00049919327,0.00019225408,0.00024917402,0.00021405278,0.0906736,0.008497596,0.02745505,0.049284283,0.817497],"study_design_scores_gemma":[0.00003637733,0.00010578025,0.001786832,0.000087455795,0.00006321104,0.00022745875,0.00009808503,0.927945,0.009728379,0.04126022,0.018636625,0.000024558187],"about_ca_topic_score_codex":0.008158749,"about_ca_topic_score_gemma":0.011062484,"teacher_disagreement_score":0.008158749,"about_ca_system_score_codex":0.0010358103,"about_ca_system_score_gemma":0.0019181508,"threshold_uncertainty_score":0.01901257},"labels":[],"label_agreement":null},{"id":"W4312438588","doi":"10.1145/3524842.3528470","title":"An empirical evaluation of GitHub copilot's code suggestions","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":281,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Correctness; Programmer; JavaScript; Programming language; Code (set theory); Source code; Java; Software engineering","score_opus":0.09023633677771457,"score_gpt":0.39519672694441305,"score_spread":0.30496039016669846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312438588","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98295045,0.0002144998,0.007387805,0.00026864998,0.000027973607,0.0003995998,0.0010967967,0.004163304,0.0034908587],"genre_scores_gemma":[0.9597022,0.00021631777,0.029495768,0.00024958336,0.000025174879,0.0005767225,0.005374507,0.0015871875,0.0027725375],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9867403,0.0054055233,0.0008039844,0.0012969032,0.0052582086,0.000494996],"domain_scores_gemma":[0.82359093,0.13514283,0.008050979,0.008560714,0.022582041,0.0020725427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009153165,0.00088223175,0.0004234974,0.002951904,0.00069755065,0.0013179886,0.0016576113,0.0010229519,0.0022333506],"category_scores_gemma":[0.1125994,0.00037294958,0.0004211164,0.0020522675,0.0011547337,0.0018761444,0.001579597,0.0011912596,0.0010761095],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0047155106,0.0041467934,0.3572329,0.0050235935,0.00037124462,0.003229229,0.04137112,0.01818536,0.041795604,0.004812973,0.046973478,0.47214225],"study_design_scores_gemma":[0.000894735,0.007980724,0.5591258,0.0013231905,0.00041011017,0.003824031,0.020531625,0.21333465,0.06486284,0.0026959213,0.1245024,0.00051391975],"about_ca_topic_score_codex":0.0035532385,"about_ca_topic_score_gemma":0.005888313,"teacher_disagreement_score":0.009153165,"about_ca_system_score_codex":0.00132462,"about_ca_system_score_gemma":0.0012340278,"threshold_uncertainty_score":0.048407197},"labels":[],"label_agreement":null},{"id":"W4312443713","doi":"10.1109/re54965.2022.00011","title":"Automated Question Answering for Improved Understanding of Compliance Requirements: A Multi-Document Study","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Jargon; Computer science; Compliance (psychology); Question answering; Natural language; Requirements engineering; Software requirements; Subject (documents); Information retrieval; Software; World Wide Web; Artificial intelligence; Software development; Software design; Linguistics","score_opus":0.13453979050935305,"score_gpt":0.380177912239094,"score_spread":0.24563812172974092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312443713","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.68814695,0.0049724127,0.2895168,0.0023975968,0.00012230023,0.0014140318,0.0029202285,0.0039764093,0.006533188],"genre_scores_gemma":[0.6974937,0.0011690018,0.28818735,0.00058363896,0.00012840987,0.0006804222,0.008430876,0.00046474853,0.002861924],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9815788,0.012153798,0.0011250318,0.0024106533,0.0023704236,0.0003611839],"domain_scores_gemma":[0.8583996,0.12375898,0.004141656,0.0042801704,0.008365941,0.0010536329],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012203212,0.0009959381,0.0009920289,0.0047966647,0.001472076,0.0028663857,0.0017558174,0.0021748866,0.002997954],"category_scores_gemma":[0.04936987,0.00044808298,0.0010756657,0.002785879,0.00093016715,0.0058416636,0.00284168,0.002092442,0.0013779209],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001249279,0.0028348737,0.038998563,0.0034886508,0.0004350696,0.0013082365,0.017098945,0.02567684,0.052799042,0.0077181393,0.014581134,0.8338112],"study_design_scores_gemma":[0.00048819583,0.0020403918,0.073717535,0.00073933834,0.0005700717,0.0030609525,0.01483421,0.749949,0.06301175,0.014113992,0.07714566,0.0003289397],"about_ca_topic_score_codex":0.005737309,"about_ca_topic_score_gemma":0.0048115286,"teacher_disagreement_score":0.012203212,"about_ca_system_score_codex":0.0015293638,"about_ca_system_score_gemma":0.0017305912,"threshold_uncertainty_score":0.064537525},"labels":[],"label_agreement":null},{"id":"W4312475219","doi":"10.1109/icsme55016.2022.00021","title":"Exploring the Notion of Risk in Code Reviewer Recommendation","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Waterloo","funders":"","keywords":"Computer science; Workload; Premise; Code review; Recommender system; Core (optical fiber); Code (set theory); Normalization (sociology); Source code; Empirical research; Set (abstract data type); Data science; Information retrieval; Risk analysis (engineering); Software","score_opus":0.11168447792625469,"score_gpt":0.300212604902923,"score_spread":0.18852812697666832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312475219","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33901432,0.0045575844,0.64601696,0.0024353035,0.00018536716,0.00024479462,0.00045255705,0.0017812916,0.00531176],"genre_scores_gemma":[0.9162028,0.00040005415,0.081848666,0.00017789291,0.00012150128,0.00006630398,0.00026556195,0.000070855785,0.00084645284],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98742336,0.0053826096,0.0009778193,0.0028067692,0.0030108332,0.0003986272],"domain_scores_gemma":[0.90266114,0.07197206,0.008738158,0.0074719912,0.007901598,0.0012550793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012154999,0.0012115372,0.0012460586,0.0030689395,0.0009058205,0.0029358072,0.0016540058,0.0018056219,0.0008465359],"category_scores_gemma":[0.086090304,0.0007621606,0.00094702333,0.002063137,0.0009892111,0.005732619,0.0015771015,0.002175347,0.0004066518],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011242337,0.0006974704,0.22723944,0.0009381069,0.0009157315,0.0005705833,0.0039023953,0.22264704,0.014866202,0.0120156435,0.005211549,0.5098715],"study_design_scores_gemma":[0.00006415003,0.00048189325,0.034368392,0.000120664605,0.00023570338,0.0006425473,0.00044465336,0.9403771,0.0053778696,0.013842216,0.003881712,0.00016301865],"about_ca_topic_score_codex":0.0058917,"about_ca_topic_score_gemma":0.0070759226,"teacher_disagreement_score":0.012154999,"about_ca_system_score_codex":0.001312049,"about_ca_system_score_gemma":0.0014505866,"threshold_uncertainty_score":0.06428254},"labels":[],"label_agreement":null},{"id":"W4312576573","doi":"10.1145/3524610.3527875","title":"Casdoc","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Code (set theory); Java; Presentation (obstetrics); Programming language; Process (computing); World Wide Web","score_opus":0.016884646640261246,"score_gpt":0.2550039565694777,"score_spread":0.23811930992921648,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312576573","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003464788,0.0005407212,0.19993697,0.00094784965,0.0013995251,0.00080974126,0.05080979,0.57705957,0.1650311],"genre_scores_gemma":[0.0431147,0.0010255812,0.2954129,0.0016121705,0.00059854885,0.0018581273,0.16437969,0.20788689,0.28411138],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982157,0.0002870944,0.00018509923,0.00042673285,0.00074842246,0.0001370123],"domain_scores_gemma":[0.98956597,0.0033431577,0.00049405603,0.003124915,0.0029202953,0.00055160344],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0021450163,0.0017497652,0.00092200225,0.0030846582,0.0011141606,0.0052400394,0.0035046022,0.0022948538,0.32088795],"category_scores_gemma":[0.015379783,0.0014320115,0.0014996864,0.0019949935,0.0006694372,0.0045950827,0.0034271555,0.0023509413,0.20129487],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005399153,0.00015057775,0.0010416175,0.0010475349,0.000037871974,0.00029584774,0.00040953883,0.0010420885,0.0057035964,0.017903995,0.77012557,0.20170183],"study_design_scores_gemma":[0.00008000492,0.000051013398,0.00044604737,0.00010580014,0.00001651028,0.0002335353,0.000043269083,0.0024106724,0.0052256323,0.003924757,0.98742515,0.000037651746],"about_ca_topic_score_codex":0.001950176,"about_ca_topic_score_gemma":0.0030307279,"teacher_disagreement_score":0.6791121,"about_ca_system_score_codex":0.0014110141,"about_ca_system_score_gemma":0.0022793142,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4312614778","doi":"10.1109/tse.2022.3222160","title":"Studying the Interplay Between the Durations and Breakages of Continuous Integration Builds","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Toronto; University of Ottawa","funders":"","keywords":"Computer science; Context (archaeology); World Wide Web; Data science; Information retrieval","score_opus":0.013719677836531744,"score_gpt":0.2563637641058254,"score_spread":0.24264408626929365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312614778","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9923495,0.00026572376,0.003710815,0.00016884842,0.000010067904,0.00005730786,0.0001362778,0.000042678617,0.0032587692],"genre_scores_gemma":[0.9955857,0.0001950517,0.0032700482,0.000043262124,0.000010835827,0.00012041439,0.00015789637,0.000022928554,0.0005938914],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9840295,0.0071409363,0.0015811999,0.0022142772,0.0040776734,0.00095646875],"domain_scores_gemma":[0.6942674,0.20598488,0.070065804,0.009284007,0.014432425,0.005965609],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014369593,0.00046721558,0.00034320707,0.0024662707,0.00091743446,0.002953742,0.00115902,0.00094532093,0.002723546],"category_scores_gemma":[0.13073301,0.00076737313,0.00035262044,0.002174781,0.0013273554,0.003908128,0.0021171707,0.0015994825,0.00039513365],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007166176,0.0015479302,0.8614068,0.00067042763,0.00029594297,0.0002444226,0.037270047,0.0021682696,0.008247212,0.003139944,0.0008379386,0.083454415],"study_design_scores_gemma":[0.000039485774,0.0012268093,0.9672998,0.00017463794,0.00015235372,0.00016307241,0.01812942,0.0027430514,0.0033984033,0.001890599,0.0046955566,0.00008677212],"about_ca_topic_score_codex":0.003868722,"about_ca_topic_score_gemma":0.00835566,"teacher_disagreement_score":0.014369593,"about_ca_system_score_codex":0.0017420199,"about_ca_system_score_gemma":0.0017793187,"threshold_uncertainty_score":0.07599455},"labels":[],"label_agreement":null},{"id":"W4312658591","doi":"10.1109/iwsc55060.2022.00020","title":"An Insight into the Reusability of Stack Overflow Code Fragments in Mobile Applications","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Codebase; Code review; Computer science; Code smell; Code reuse; Source code; Code (set theory); Software quality; Java; Static program analysis; Reuse; Software; Software engineering; Usability; Software bug; World Wide Web; Software development; Programming language; Operating system; Engineering","score_opus":0.014727693863767011,"score_gpt":0.2944479045720571,"score_spread":0.2797202107082901,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312658591","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98810565,0.00045779863,0.009052366,0.00012339569,0.000014713849,0.000073606934,0.00046886285,0.00071372883,0.0009898674],"genre_scores_gemma":[0.9840217,0.00016138775,0.01335495,0.000052922074,0.000017623397,0.000063425934,0.0010657272,0.00021400639,0.0010482065],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99727625,0.0005953562,0.00023914466,0.0005519275,0.0011250793,0.00021220166],"domain_scores_gemma":[0.9764568,0.013257596,0.0043402836,0.0018004965,0.0036478715,0.00049686286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026987586,0.0005175094,0.00037151697,0.004889947,0.00063335686,0.0010684752,0.0005280161,0.0006771642,0.0009677163],"category_scores_gemma":[0.025451709,0.00028251525,0.0005391044,0.0019676043,0.00060199574,0.0017478025,0.0011864428,0.00057790196,0.00031899975],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008123285,0.00044396127,0.63483036,0.0009142693,0.0002525692,0.0030241464,0.017788297,0.005973856,0.052303374,0.0019871402,0.0035581824,0.27811146],"study_design_scores_gemma":[0.000037724465,0.00088152685,0.85398906,0.00025796177,0.00027445354,0.003073765,0.0056606834,0.099828355,0.022956625,0.0033876677,0.009503433,0.00014873676],"about_ca_topic_score_codex":0.0036383073,"about_ca_topic_score_gemma":0.0067440984,"teacher_disagreement_score":0.004889947,"about_ca_system_score_codex":0.00040955926,"about_ca_system_score_gemma":0.00053261727,"threshold_uncertainty_score":0.014272571},"labels":[],"label_agreement":null},{"id":"W4312691020","doi":"10.1145/3524842.3528527","title":"Refactoring debt","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Technical debt; Codebase; Computer science; Java; Software engineering; Source code; Code (set theory); Timeline; Programming language; Software development; Software; Set (abstract data type)","score_opus":0.022858249623685687,"score_gpt":0.26369360928006147,"score_spread":0.24083535965637579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312691020","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5818684,0.002956816,0.34731635,0.003974172,0.0007216198,0.0016343866,0.0017794154,0.015582998,0.044165764],"genre_scores_gemma":[0.6802159,0.0014594367,0.27869737,0.0015949759,0.00018607533,0.00050197286,0.0030638697,0.0052713803,0.029009104],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9836361,0.004310658,0.0021171768,0.0020057685,0.0069527947,0.0009774562],"domain_scores_gemma":[0.8755214,0.035279218,0.020136913,0.041662622,0.02556375,0.0018360745],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017577369,0.0011696339,0.0006140138,0.0048453314,0.001524798,0.0029389106,0.0026263376,0.0015413587,0.0037653171],"category_scores_gemma":[0.11522565,0.0009692112,0.000808575,0.0033529927,0.001307231,0.0045538694,0.0042214347,0.0025833775,0.0020650947],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002702319,0.00048725688,0.15651828,0.0012938065,0.00023134206,0.0030291593,0.0141064,0.005370378,0.03228672,0.02372473,0.027863223,0.73481846],"study_design_scores_gemma":[0.00019454907,0.00093924796,0.1455532,0.0027500032,0.0005587029,0.009394742,0.010537779,0.04937618,0.06297631,0.052467857,0.6647661,0.0004853998],"about_ca_topic_score_codex":0.0055218027,"about_ca_topic_score_gemma":0.006671386,"teacher_disagreement_score":0.017577369,"about_ca_system_score_codex":0.0019660578,"about_ca_system_score_gemma":0.0052065244,"threshold_uncertainty_score":0.092959106},"labels":[],"label_agreement":null},{"id":"W4312695097","doi":"10.1145/3524842.3527932","title":"Code review practices for refactoring changes","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Code review; Software engineering; Source code; Code (set theory); Process (computing); Coding (social sciences); Software; Programming language; Software quality; Software development","score_opus":0.12551164725424324,"score_gpt":0.38498876877203725,"score_spread":0.259477121517794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312695097","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6627567,0.022521809,0.22367829,0.020283706,0.0014943911,0.0064574415,0.0008628421,0.0041256333,0.057819188],"genre_scores_gemma":[0.8351644,0.006385294,0.13551092,0.0026300447,0.00033018316,0.0028654204,0.00092858175,0.001576821,0.014608369],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8463538,0.081006415,0.013997031,0.012827718,0.042716913,0.0030981384],"domain_scores_gemma":[0.5245782,0.23008627,0.05764811,0.02587961,0.15343554,0.008372345],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07886985,0.00080338126,0.00090018637,0.011242382,0.006985874,0.005005829,0.0030827231,0.0021300563,0.0029113921],"category_scores_gemma":[0.3226168,0.00073665986,0.0010994257,0.005475033,0.0027498102,0.0062373863,0.006308476,0.0026174719,0.0018524006],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035312018,0.00028887365,0.06458561,0.0048452034,0.0002420305,0.0020635864,0.23833494,0.0014788052,0.027864296,0.0070429486,0.022016143,0.63088447],"study_design_scores_gemma":[0.00014951902,0.0010970372,0.16338658,0.011728411,0.00045812764,0.0076053236,0.14428829,0.011169628,0.029594373,0.015274732,0.6141917,0.0010562476],"about_ca_topic_score_codex":0.005931735,"about_ca_topic_score_gemma":0.010177571,"teacher_disagreement_score":0.9211302,"about_ca_system_score_codex":0.005697119,"about_ca_system_score_gemma":0.01561462,"threshold_uncertainty_score":0.4171086},"labels":[],"label_agreement":null},{"id":"W4312753189","doi":"10.1109/icsme55016.2022.00075","title":"Mining Annotation Usage Rules: A Case Study with MicroProfile","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); University of Alberta","funders":"","keywords":"Computer science; Java; Application programming interface; Annotation; Microservices; Documentation; Process (computing); Reuse; Software engineering; Data mining; Programming language; Information retrieval; Artificial intelligence","score_opus":0.0223596113575428,"score_gpt":0.27618583107003997,"score_spread":0.2538262197124972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312753189","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9677844,0.00040205303,0.02694322,0.0008426879,0.000021157055,0.0002653111,0.001088201,0.0007114578,0.001941562],"genre_scores_gemma":[0.87508065,0.00048854266,0.11850231,0.00027024047,0.000021725775,0.00026213023,0.002786776,0.00039038673,0.0021972],"study_design_codex":"observational","study_design_gemma":"case_report","domain_scores_codex":[0.98829764,0.004577315,0.0013911991,0.0015497392,0.003632936,0.00055128377],"domain_scores_gemma":[0.88254344,0.09426202,0.0075486344,0.0074021574,0.007173425,0.0010703024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008795461,0.00067749026,0.0006696306,0.003677146,0.0020790445,0.0016051565,0.0023700646,0.0024307382,0.0007794099],"category_scores_gemma":[0.04583298,0.0006564463,0.00083408906,0.0042736,0.0013664719,0.0026160867,0.0015948564,0.0016880576,0.00034474683],"study_design_candidate":"case_report","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010348284,0.0034780472,0.4790304,0.002808539,0.0003920868,0.04759726,0.049151193,0.028464932,0.027223364,0.0060117794,0.011541631,0.3432659],"study_design_scores_gemma":[0.00047761973,0.0025228558,0.3459736,0.0018722311,0.0008333701,0.04444176,0.063788764,0.29424453,0.12939498,0.013919282,0.101893134,0.0006377866],"about_ca_topic_score_codex":0.01008882,"about_ca_topic_score_gemma":0.021806942,"teacher_disagreement_score":0.01008882,"about_ca_system_score_codex":0.0010114792,"about_ca_system_score_gemma":0.0015112587,"threshold_uncertainty_score":0.046515346},"labels":[],"label_agreement":null},{"id":"W4312757284","doi":"10.1109/icsme55016.2022.00012","title":"An Empirical Study on Performance Bugs in Deep Learning Frameworks","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software bug; Artificial intelligence; Empirical research; Machine learning; Deep learning; Performance improvement; Quality (philosophy); Software; Software engineering; Programming language","score_opus":0.022385434239069864,"score_gpt":0.32553213727032954,"score_spread":0.3031467030312597,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312757284","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9912683,0.0011975651,0.0033820346,0.0005497523,0.000033888566,0.000104860854,0.0012160209,0.0010552803,0.0011923582],"genre_scores_gemma":[0.990624,0.00043694535,0.0052862023,0.000171477,0.000026994192,0.000115169336,0.0023957305,0.00031927388,0.0006241006],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9797641,0.005024857,0.0026188218,0.0031005016,0.008143401,0.0013482879],"domain_scores_gemma":[0.72760296,0.15922864,0.06361702,0.013428786,0.031793512,0.004329027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013326056,0.0009772303,0.00059721584,0.008215399,0.0009458529,0.001605855,0.001769827,0.0012393921,0.0013498309],"category_scores_gemma":[0.15958975,0.0007299616,0.0007509617,0.0052182665,0.0019509047,0.0043720626,0.0017880806,0.0023222165,0.00047716906],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004939637,0.00089972024,0.87523204,0.0009797128,0.0002199594,0.00091123354,0.004684616,0.0038433978,0.0025965485,0.0015312554,0.008460227,0.10014737],"study_design_scores_gemma":[0.00011235844,0.0014096757,0.9113221,0.0009130083,0.00032810168,0.0034362162,0.0065544727,0.050309338,0.008496424,0.0025411916,0.014359378,0.0002177714],"about_ca_topic_score_codex":0.0054808883,"about_ca_topic_score_gemma":0.006651284,"teacher_disagreement_score":0.013326056,"about_ca_system_score_codex":0.0014489638,"about_ca_system_score_gemma":0.0017209166,"threshold_uncertainty_score":0.07047582},"labels":[],"label_agreement":null},{"id":"W4312844600","doi":"10.1007/978-3-031-11686-5_3","title":"Semantic History Slicing","year":2022,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of British Columbia","funders":"","keywords":"Slicing; Program slicing; Computer science; Feature (linguistics); Dependency (UML); Key (lock); Semantics (computer science); Programming language; Information retrieval; Artificial intelligence; World Wide Web; Computer security","score_opus":0.0298441271368807,"score_gpt":0.22801552053454346,"score_spread":0.19817139339766277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312844600","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0044480637,0.003203742,0.37313235,0.0010746877,0.0007116389,0.00008330611,0.00076318625,0.0047541987,0.6118288],"genre_scores_gemma":[0.16534634,0.009127771,0.28457826,0.0009880869,0.0008872553,0.00013777331,0.005418611,0.0056841443,0.5278318],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99960965,0.000043333126,0.000021183449,0.00009001535,0.00019897606,0.00003677908],"domain_scores_gemma":[0.99936086,0.00017641186,0.000020516569,0.00027832494,0.00013715537,0.000026841772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00044497382,0.00076632167,0.00045088324,0.0018311231,0.00093953294,0.0019906245,0.00092155737,0.00047160688,0.038731933],"category_scores_gemma":[0.0014297074,0.00052677747,0.0006127628,0.0019443849,0.0019288126,0.005163475,0.0015554233,0.0018881739,0.009391206],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003457971,0.000015252271,0.00015049806,0.00014526468,0.000008451194,0.000087691435,0.0003771109,0.0020586387,0.0022633227,0.66244775,0.05241817,0.2799933],"study_design_scores_gemma":[0.0000073034244,0.000011957362,0.00016223002,0.00013149594,0.000020888307,0.00023515981,0.00014898623,0.006385905,0.005686099,0.43972665,0.5474682,0.0000151185795],"about_ca_topic_score_codex":0.0033931227,"about_ca_topic_score_gemma":0.005802937,"teacher_disagreement_score":0.038731933,"about_ca_system_score_codex":0.0013002075,"about_ca_system_score_gemma":0.0013962121,"threshold_uncertainty_score":0.1295712},"labels":[],"label_agreement":null},{"id":"W4312902756","doi":"10.1145/3524842.3528034","title":"Is refactoring always a good egg?","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nemzeti Kutatási, Fejlesztési és Innovaciós Alap; Wageningen University and Research; Kindai University; Innopolis University; Szegedi Tudományegyetem; Universität Stuttgart; Innovációs és Technológiai Minisztérium; Magyar Tudományos Akadémia; Kanazawa University; Massey University; University of Alberta; Nanzan University","keywords":"Code refactoring; Maintainability; Computer science; Commit; Code smell; Programming language; Code (set theory); Software quality; Software maintenance; Software engineering; Software; Software system; Software development; Database","score_opus":0.033026876320357446,"score_gpt":0.2793055967836429,"score_spread":0.24627872046328547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312902756","genre_codex":"empirical","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86993253,0.012704697,0.02438973,0.012840189,0.0012110865,0.00014526643,0.056704223,0.0052232966,0.016848989],"genre_scores_gemma":[0.8392366,0.0032360062,0.033262797,0.0036118147,0.0005086022,0.00011945816,0.110136695,0.0015818446,0.008306101],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934755,0.0011455945,0.0005877578,0.002138555,0.002162537,0.00049002946],"domain_scores_gemma":[0.9478245,0.022121044,0.008656977,0.014417938,0.005526934,0.0014525823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056884605,0.00055627624,0.00072567625,0.0032524217,0.0010694788,0.0017678814,0.0016076595,0.0015172081,0.0028034986],"category_scores_gemma":[0.047853608,0.00048614142,0.0006819509,0.0044887764,0.0009755744,0.0046860552,0.0011935214,0.0015970134,0.0020585544],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072572666,0.00024469892,0.53366417,0.0013909629,0.00037636454,0.0013614204,0.0018654936,0.0028752335,0.005581211,0.006187814,0.14298253,0.30274445],"study_design_scores_gemma":[0.00014793081,0.00052018464,0.57328415,0.0008238055,0.0002739031,0.0055427365,0.0034438253,0.014787111,0.008343501,0.016264528,0.37637565,0.00019275257],"about_ca_topic_score_codex":0.007023103,"about_ca_topic_score_gemma":0.011788809,"teacher_disagreement_score":0.007023103,"about_ca_system_score_codex":0.00075459474,"about_ca_system_score_gemma":0.0010245094,"threshold_uncertainty_score":0.030083776},"labels":[],"label_agreement":null},{"id":"W4312943943","doi":"10.1007/978-3-031-21037-2_10","title":"Novice Type Error Diagnosis with Natural Language Models","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Flexibility (engineering); Metric (unit); Constraint (computer-aided design); Natural language; Type (biology); Artificial intelligence; Language model; State (computer science); Data type; Machine learning; Natural language understanding; Natural language processing; Algorithm; Programming language","score_opus":0.01822127109074478,"score_gpt":0.26224330893068043,"score_spread":0.24402203783993565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312943943","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013942597,0.00024873315,0.9667691,0.00046886405,0.0000861419,0.000090446825,0.00032603816,0.010348559,0.007719535],"genre_scores_gemma":[0.25238273,0.00034943287,0.7312351,0.0004673216,0.000086449094,0.00011286556,0.0013148064,0.0012747212,0.012776641],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99806434,0.000519927,0.00011260462,0.00044312255,0.00073529955,0.00012478467],"domain_scores_gemma":[0.99204093,0.004760431,0.00041226414,0.001655001,0.0009984712,0.00013292265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014111894,0.0009862403,0.0005987481,0.0011365694,0.00037798152,0.0016391089,0.0024069885,0.0014532096,0.007240238],"category_scores_gemma":[0.01137854,0.0005840718,0.001149443,0.0004560065,0.0007763274,0.0037869215,0.0016981895,0.002259471,0.0031089138],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004489947,0.0005263388,0.0052617774,0.00049193256,0.000099992205,0.0015556981,0.00076927524,0.06356062,0.019278612,0.050270207,0.033261936,0.82447463],"study_design_scores_gemma":[0.000055622422,0.00017804728,0.00095004606,0.00019334185,0.0000821391,0.0017680641,0.00024385615,0.76696324,0.0478326,0.15373746,0.0279296,0.00006605346],"about_ca_topic_score_codex":0.0019336436,"about_ca_topic_score_gemma":0.0039275107,"teacher_disagreement_score":0.007240238,"about_ca_system_score_codex":0.0005994661,"about_ca_system_score_gemma":0.0010863122,"threshold_uncertainty_score":0.024221003},"labels":[],"label_agreement":null},{"id":"W4312945154","doi":"10.1145/3524610.3527924","title":"LAMNER","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Programming language; Code (set theory); Abstract syntax tree; Semantics (computer science); Security token; Source code; Code generation; Theoretical computer science; Parsing","score_opus":0.015321893644444919,"score_gpt":0.24837919431745575,"score_spread":0.23305730067301084,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312945154","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008561998,0.00226341,0.20991385,0.0067604943,0.0036318041,0.0005419157,0.014538935,0.069627486,0.6841602],"genre_scores_gemma":[0.07412875,0.0019606787,0.10793033,0.0028963028,0.0006557794,0.000511956,0.03282247,0.013588732,0.7655051],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998408,0.00025476862,0.00010778193,0.0004667632,0.0006050804,0.00015760532],"domain_scores_gemma":[0.99735475,0.00041681965,0.00011931103,0.00085316267,0.0009426914,0.0003132498],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013419228,0.0010672489,0.00063769764,0.0020651694,0.0015949408,0.003546237,0.0019232387,0.0019763862,0.3133181],"category_scores_gemma":[0.0050250418,0.00059729384,0.0008255497,0.001249935,0.0006465997,0.003887991,0.0033172464,0.0018446846,0.27481467],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003392614,0.000119987555,0.0014913756,0.00050741696,0.00002895533,0.0005336476,0.00041471393,0.00180531,0.007957395,0.052392792,0.44274932,0.49165985],"study_design_scores_gemma":[0.0000285142,0.000030220459,0.0004477035,0.00010199883,0.000012106787,0.00038478355,0.00011735459,0.0038948248,0.005389379,0.012955638,0.9766028,0.000034619556],"about_ca_topic_score_codex":0.0028142757,"about_ca_topic_score_gemma":0.0036813673,"teacher_disagreement_score":0.3133181,"about_ca_system_score_codex":0.0013254654,"about_ca_system_score_gemma":0.001973886,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4312969456","doi":"10.1109/icsme55016.2022.00035","title":"Stronger Together: On Combining Relationships in Architectural Recovery Approaches","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"IBM Canada","keywords":"Computer science; Representation (politics); Process (computing); Architecture; Similarity (geometry); Software architecture; Class (philosophy); Data mining; Software engineering; Software; Artificial intelligence; Programming language","score_opus":0.0853003857517976,"score_gpt":0.25245563229853235,"score_spread":0.16715524654673475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312969456","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01446468,0.0012403716,0.96917087,0.002462512,0.00007215217,0.00027404987,0.00008568805,0.00072654884,0.011503026],"genre_scores_gemma":[0.21150818,0.0016440718,0.77933306,0.001166976,0.00017942196,0.0004179486,0.00058852276,0.000837244,0.0043245372],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.95171875,0.024567932,0.0026060408,0.005250952,0.014144975,0.0017113928],"domain_scores_gemma":[0.9317127,0.03268821,0.0072180796,0.018767541,0.00819739,0.0014161501],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03555468,0.002747932,0.0016737228,0.012213016,0.0043340838,0.009083548,0.005806417,0.0037816635,0.0069140196],"category_scores_gemma":[0.086113475,0.0019332186,0.0029698675,0.008684771,0.005953416,0.038359474,0.019530235,0.009071092,0.0020467786],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025808968,0.00038792842,0.010583373,0.0007722938,0.00043227908,0.000658964,0.015172484,0.026874393,0.0049310466,0.4372224,0.005861327,0.49684542],"study_design_scores_gemma":[0.00006937006,0.00032431862,0.0039676633,0.0010536257,0.0006086332,0.0009430787,0.008424662,0.17471895,0.006395034,0.7340095,0.06928879,0.00019639492],"about_ca_topic_score_codex":0.0035296066,"about_ca_topic_score_gemma":0.0055641625,"teacher_disagreement_score":0.03555468,"about_ca_system_score_codex":0.0026582857,"about_ca_system_score_gemma":0.0029034452,"threshold_uncertainty_score":0.18803334},"labels":[],"label_agreement":null},{"id":"W4312986843","doi":"10.2139/ssrn.4237049","title":"Detecting Unexpected Behaviors in Scenario-Based Specification: Asystematic Literature Review","year":2022,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Psychology","score_opus":0.010678848907179923,"score_gpt":0.2603485214004621,"score_spread":0.24966967249328217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312986843","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060880963,0.5067576,0.4028147,0.007423175,0.0011040344,0.000662886,0.0015206742,0.0015453198,0.017290667],"genre_scores_gemma":[0.43294567,0.3416642,0.2148471,0.0019627397,0.0014139122,0.00044380786,0.0038482936,0.00023496668,0.0026393323],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9865637,0.004717495,0.0016651166,0.0026911101,0.0038621894,0.0005003455],"domain_scores_gemma":[0.872463,0.10301841,0.005864733,0.0061236094,0.011843393,0.000686746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010839799,0.0024212983,0.0016761842,0.009823337,0.0008889917,0.005187687,0.0050612064,0.0026998702,0.003587386],"category_scores_gemma":[0.063978225,0.0009839339,0.0018648676,0.007309003,0.0022203117,0.008684668,0.00297753,0.0020118752,0.000989426],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022376758,0.0003817994,0.027592398,0.016210511,0.0007371655,0.0005525115,0.0012187023,0.021603404,0.0018802147,0.028738406,0.0050020805,0.8958591],"study_design_scores_gemma":[0.00014392621,0.0019283274,0.051048983,0.054370984,0.0043012854,0.0087335035,0.014006162,0.38694465,0.031035518,0.18632472,0.26045737,0.00070454006],"about_ca_topic_score_codex":0.0044938913,"about_ca_topic_score_gemma":0.0037276843,"teacher_disagreement_score":0.010839799,"about_ca_system_score_codex":0.0019684713,"about_ca_system_score_gemma":0.0050977645,"threshold_uncertainty_score":0.057327032},"labels":[],"label_agreement":null},{"id":"W4313137049","doi":"10.1145/3524610.3527921","title":"An exploratory study on code attention in BERT","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Automatic summarization; Transformer; Artificial intelligence; Programming language; Natural language processing; Natural language; Source code; Engineering","score_opus":0.04139991490038435,"score_gpt":0.3098199987926032,"score_spread":0.26842008389221883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313137049","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9440353,0.00097341835,0.005645834,0.0017275246,0.00002869074,0.000065126645,0.0003002069,0.0001852916,0.047038548],"genre_scores_gemma":[0.9926865,0.00021714647,0.0012929857,0.00032473524,0.00001770366,0.000028379327,0.0002931468,0.000084614796,0.0050548096],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99905604,0.00030337484,0.000026927666,0.00018053819,0.0002911976,0.00014189185],"domain_scores_gemma":[0.9891083,0.00786043,0.00082825456,0.0005400022,0.0011402243,0.00052277284],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012395851,0.00025360717,0.00024536785,0.0013955262,0.0012227905,0.0020327063,0.0007572312,0.0009186968,0.0095440885],"category_scores_gemma":[0.024866095,0.00021801286,0.0001660141,0.0023648052,0.0013341467,0.0041806833,0.0014046277,0.0011502496,0.0009827965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009790327,0.00092888135,0.38644183,0.0009445043,0.00009639982,0.0023134102,0.117474005,0.005664357,0.010844645,0.07331169,0.024030652,0.3769706],"study_design_scores_gemma":[0.000072564464,0.00067522516,0.586931,0.00055494165,0.00011902654,0.0026293292,0.09228006,0.06545902,0.0069023934,0.056275558,0.18795636,0.00014443132],"about_ca_topic_score_codex":0.018199734,"about_ca_topic_score_gemma":0.017558826,"teacher_disagreement_score":0.018199734,"about_ca_system_score_codex":0.0018796247,"about_ca_system_score_gemma":0.0008992985,"threshold_uncertainty_score":0.03618759},"labels":[],"label_agreement":null},{"id":"W4313142332","doi":"10.1109/tse.2022.3220740","title":"A Comprehensive Investigation of the Impact of Class Overlap on Software Defect Prediction","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Science Foundation of Jiangsu Province; National Natural Science Foundation of China","keywords":"Computer science; Class (philosophy); Software; Data mining; Rank (graph theory); Feature (linguistics); Machine learning; Identification (biology); Artificial intelligence; Mathematics","score_opus":0.018022052625698195,"score_gpt":0.24269995323414043,"score_spread":0.22467790060844223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313142332","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9494919,0.0014135552,0.044702947,0.0003754034,0.000060106377,0.0000690631,0.0011643672,0.0010114581,0.0017112552],"genre_scores_gemma":[0.9770493,0.00025336634,0.019386556,0.00007075118,0.000030875788,0.00004235473,0.0027446705,0.00006720523,0.0003549891],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99612576,0.0008478395,0.00028111125,0.0011709228,0.0012792719,0.00029519826],"domain_scores_gemma":[0.9789549,0.013159288,0.0021362759,0.0028644467,0.0023116036,0.00057336927],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053457017,0.0010747955,0.0008930256,0.0026322287,0.00071458315,0.0013672692,0.0011623633,0.00089289626,0.00051210314],"category_scores_gemma":[0.020804431,0.00029112792,0.001038219,0.002044341,0.00068627344,0.0026446395,0.0013548855,0.0012556596,0.00026221533],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063045934,0.00066101947,0.4232097,0.00042778772,0.00055296835,0.0007890104,0.00064355525,0.25902307,0.008562104,0.0017360365,0.007582851,0.29618144],"study_design_scores_gemma":[0.000021394577,0.00040812956,0.102615416,0.00007691019,0.00010890874,0.0005580258,0.00050228235,0.88276607,0.007388038,0.0026023833,0.002906167,0.000046248177],"about_ca_topic_score_codex":0.006285302,"about_ca_topic_score_gemma":0.006757194,"teacher_disagreement_score":0.006285302,"about_ca_system_score_codex":0.00073660014,"about_ca_system_score_gemma":0.0009705519,"threshold_uncertainty_score":0.028271139},"labels":[],"label_agreement":null},{"id":"W4313192501","doi":"10.1017/s0890060422000166","title":"Machine learning in requirements elicitation: a literature review","year":2022,"lang":"en","type":"review","venue":"Artificial intelligence for engineering design analysis and manufacturing","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Requirements elicitation; Computer science; Preprocessor; Construct (python library); Process (computing); Expert elicitation; Data pre-processing; Requirements management; Artificial intelligence; Machine learning; Information retrieval; Requirements analysis; Software","score_opus":0.09816058246718821,"score_gpt":0.34776332462493664,"score_spread":0.24960274215774841,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313192501","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014130815,0.9908806,0.003440715,0.0013289182,0.00013824046,0.000085383064,0.00013271424,0.000040560688,0.0025397365],"genre_scores_gemma":[0.013304039,0.9810781,0.004376479,0.00047398754,0.00014361196,0.00012925503,0.0002306656,0.000013543271,0.0002502698],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.99483913,0.0022830726,0.00094871176,0.00041352303,0.0013408737,0.00017473525],"domain_scores_gemma":[0.9330817,0.058064915,0.002648571,0.00071581046,0.0052313535,0.0002576431],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.008232978,0.0012543012,0.001769774,0.01084163,0.0006365132,0.002249028,0.0020887773,0.001682233,0.0033123596],"category_scores_gemma":[0.023726543,0.00071501266,0.0018445388,0.015034048,0.0008500622,0.0032187751,0.0012019561,0.0011074202,0.00095039123],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010913724,0.00012504583,0.0010979357,0.16840921,0.0003651976,0.0003267821,0.00065737666,0.0024543263,0.000717588,0.0063976403,0.011596491,0.8077433],"study_design_scores_gemma":[0.000071349066,0.0003973119,0.009180802,0.32724968,0.0030430513,0.0021057706,0.0028020765,0.0066071996,0.0036417362,0.01357057,0.6311354,0.00019501371],"about_ca_topic_score_codex":0.004097179,"about_ca_topic_score_gemma":0.004326658,"teacher_disagreement_score":0.9891584,"about_ca_system_score_codex":0.0022445717,"about_ca_system_score_gemma":0.005954712,"threshold_uncertainty_score":0.043540657},"labels":[],"label_agreement":null},{"id":"W4313547553","doi":"10.1145/3551349.3559568","title":"Extraction and Management of Rationale","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Structuring; Consistency (knowledge bases); Exploit; Information extraction; Process (computing); Coherence (philosophical gambling strategy); Software engineering; Natural language; Artificial intelligence; Data science; Knowledge management; Programming language; Computer security","score_opus":0.021474954584098038,"score_gpt":0.2750518449066414,"score_spread":0.2535768903225434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313547553","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011357848,0.0017809342,0.935622,0.0036714748,0.0005548572,0.0027926636,0.0101208985,0.011708363,0.022390941],"genre_scores_gemma":[0.05059947,0.0014551322,0.9162342,0.0004859406,0.00023857562,0.0011708662,0.020511841,0.0014950543,0.0078089065],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97284985,0.00887949,0.0037130108,0.0025825356,0.011307845,0.0006673407],"domain_scores_gemma":[0.9232891,0.036040004,0.0058711274,0.01540522,0.018412832,0.000981693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021571921,0.0017900799,0.001550263,0.019988874,0.0023988045,0.009176518,0.0029691742,0.0020390442,0.008726559],"category_scores_gemma":[0.086118,0.0015868163,0.0022249392,0.008408015,0.0015452647,0.008425445,0.0071993372,0.0044016927,0.006309678],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018647607,0.00029222626,0.0037262223,0.0017606885,0.00012932289,0.0010540446,0.0026667418,0.0035346719,0.013168461,0.0711003,0.05514646,0.84723437],"study_design_scores_gemma":[0.00017867006,0.00023309859,0.0055121337,0.0024077527,0.00034654597,0.0015291534,0.0028291661,0.08015841,0.045200713,0.21634266,0.644926,0.0003357439],"about_ca_topic_score_codex":0.0025393327,"about_ca_topic_score_gemma":0.0037536821,"teacher_disagreement_score":0.021571921,"about_ca_system_score_codex":0.0018888923,"about_ca_system_score_gemma":0.009001671,"threshold_uncertainty_score":0.1140846},"labels":[],"label_agreement":null},{"id":"W4313563522","doi":"10.1145/3551349.3556896","title":"Not All Dependencies are Equal: An Empirical Study on Production Dependencies in NPM","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Production (economics); Software security assurance; Software; Software engineering; Code (set theory); Empirical research; Backporting; Software development; Computer security; Software construction; Database; Programming language; Information security; Security service","score_opus":0.10501677049643932,"score_gpt":0.35739637521049467,"score_spread":0.25237960471405535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313563522","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9916414,0.00019431989,0.0011146914,0.0007677743,0.000010038275,0.000060214377,0.00013327712,0.00001044915,0.0060678837],"genre_scores_gemma":[0.9980628,0.00010489184,0.0007277133,0.00013177279,0.000011222862,0.00006467476,0.00023731499,0.000024861234,0.0006347234],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9917612,0.0046087047,0.00050057855,0.0007166153,0.0016591386,0.00075382856],"domain_scores_gemma":[0.7514598,0.19162656,0.032770943,0.00836347,0.00824111,0.0075381254],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012179524,0.00027223272,0.00045017994,0.0013351649,0.0026233518,0.0021533812,0.0018824027,0.0011602544,0.008269885],"category_scores_gemma":[0.11117197,0.0003846278,0.00040399056,0.0022233434,0.0024230129,0.005788053,0.0029430212,0.0040847184,0.000989035],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002746193,0.0013155364,0.9466726,0.00013948652,0.00007831741,0.0007542383,0.021867823,0.0005993448,0.00018137954,0.005025175,0.0028757947,0.020215614],"study_design_scores_gemma":[0.000044267483,0.00034741455,0.9159928,0.00029318975,0.000061279985,0.0006841772,0.05886351,0.0060402844,0.00018692843,0.005645345,0.011809606,0.000031133593],"about_ca_topic_score_codex":0.008442421,"about_ca_topic_score_gemma":0.0068755676,"teacher_disagreement_score":0.012179524,"about_ca_system_score_codex":0.001864114,"about_ca_system_score_gemma":0.0020641005,"threshold_uncertainty_score":0.064412236},"labels":[],"label_agreement":null},{"id":"W4313563578","doi":"10.1145/3551349.3559537","title":"AntiCopyPaster: Extracting Code Duplicates As Soon As They Are Introduced in the IDE","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Code refactoring; Computer science; Plug-in; Programming language; Code (set theory); Workflow; Source code; Fragment (logic); Software engineering; Database; Software; Set (abstract data type)","score_opus":0.0228698490153916,"score_gpt":0.29168186762216863,"score_spread":0.26881201860677706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313563578","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04683753,0.0008045135,0.62385225,0.00048517567,0.00033786162,0.0007628074,0.004546067,0.3169559,0.0054179356],"genre_scores_gemma":[0.11967036,0.00040463023,0.83618367,0.00039236606,0.00010937553,0.00048166234,0.010815948,0.021609064,0.010332886],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9944317,0.00052420597,0.00051520334,0.001621304,0.0025664282,0.00034120577],"domain_scores_gemma":[0.9773006,0.009962376,0.0031920946,0.005646006,0.0033527876,0.00054623064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053015556,0.0023221283,0.0014981854,0.0045524864,0.0010452913,0.0030967603,0.0032254204,0.0017182219,0.00403868],"category_scores_gemma":[0.028033454,0.0017842556,0.0015105068,0.0015689342,0.0010095119,0.0043064402,0.002558846,0.0027562238,0.003771205],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010078603,0.00050800387,0.04198181,0.0018906295,0.00024470763,0.0010431956,0.0025718582,0.008102429,0.048634045,0.009008827,0.06920424,0.8158024],"study_design_scores_gemma":[0.00029459788,0.000797624,0.026527723,0.0007028684,0.00046359395,0.002661186,0.0007328881,0.4021716,0.27912256,0.01650941,0.2694995,0.0005165204],"about_ca_topic_score_codex":0.0040748897,"about_ca_topic_score_gemma":0.008107359,"teacher_disagreement_score":0.0053015556,"about_ca_system_score_codex":0.00089943095,"about_ca_system_score_gemma":0.0036159891,"threshold_uncertainty_score":0.028037608},"labels":[],"label_agreement":null},{"id":"W4313563632","doi":"10.1145/3551349.3556931","title":"How Useful is Code Change Information for Fault Localization in Continuous Integration?","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Merge (version control); Code (set theory); Fault (geology); Software; Legacy system; Data mining; Process (computing); Real-time computing; Information retrieval; Programming language","score_opus":0.03264587802112539,"score_gpt":0.26864196524963474,"score_spread":0.23599608722850934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313563632","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6889488,0.023677425,0.1865568,0.003425626,0.00087976933,0.00050192996,0.015894206,0.07130435,0.00881104],"genre_scores_gemma":[0.8245868,0.0018714707,0.13905278,0.00071727706,0.00023052975,0.00018472076,0.030023972,0.0014912394,0.0018411916],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99449277,0.0005709597,0.00043533114,0.0020672944,0.0020512654,0.00038236947],"domain_scores_gemma":[0.9852319,0.0050144037,0.0022454413,0.003753243,0.003324775,0.0004304072],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022984203,0.0020980062,0.0016569574,0.008211809,0.0008890018,0.002258483,0.002871011,0.0021177812,0.0008512618],"category_scores_gemma":[0.02222277,0.0005085025,0.0012207715,0.0050345752,0.0011987152,0.005243881,0.0020558997,0.001895054,0.0011509944],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009956055,0.000539318,0.17557433,0.0015913493,0.00042814715,0.0006956882,0.00091929245,0.030215586,0.029197222,0.001564794,0.041611068,0.71666753],"study_design_scores_gemma":[0.0004303249,0.00187213,0.20813532,0.0006804396,0.0010746837,0.003693773,0.0017973107,0.58294666,0.10508036,0.012759411,0.08118558,0.00034410006],"about_ca_topic_score_codex":0.009549797,"about_ca_topic_score_gemma":0.014500887,"teacher_disagreement_score":0.009549797,"about_ca_system_score_codex":0.0010774394,"about_ca_system_score_gemma":0.0012584529,"threshold_uncertainty_score":0.01898843},"labels":[],"label_agreement":null},{"id":"W4313563813","doi":"10.1145/3551349.3559548","title":"Automatic Code Documentation Generation Using GPT-3","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":90,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates; University of Calgary","keywords":"Documentation; Computer science; Artifact (error); Automation; Software documentation; Source code; Internal documentation; Code (set theory); Programming language; Technical documentation; Software engineering; Software; Artificial intelligence; Software development; Software development process; Software construction; Set (abstract data type); Engineering","score_opus":0.04663174896728765,"score_gpt":0.31805628152906923,"score_spread":0.27142453256178156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313563813","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09376337,0.0018013581,0.70206475,0.0012222559,0.0008510767,0.0006669699,0.006678779,0.18185389,0.011097603],"genre_scores_gemma":[0.32489446,0.0007074372,0.62488997,0.00082498917,0.00009820449,0.00072449126,0.030302849,0.0075052925,0.010052261],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99886715,0.00029753335,0.00007503847,0.0003676966,0.00031041243,0.000082231025],"domain_scores_gemma":[0.99578637,0.0018476625,0.0002514501,0.00082941476,0.0011557564,0.00012934695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014816668,0.0012383029,0.00064581295,0.0017815102,0.0005902523,0.0013141887,0.0017574481,0.001540763,0.0043990603],"category_scores_gemma":[0.011088588,0.0005989439,0.001174429,0.0011756854,0.0005713707,0.0017805692,0.0016167107,0.002528493,0.0041236044],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031840024,0.0003137702,0.0051749474,0.001117698,0.00012830227,0.00082602655,0.000651932,0.065426916,0.02626764,0.005394862,0.1085606,0.78581905],"study_design_scores_gemma":[0.000102861224,0.00015779374,0.0015752213,0.00016329931,0.000052856903,0.0006295107,0.00013555067,0.9168251,0.031411078,0.008822433,0.04006518,0.00005905528],"about_ca_topic_score_codex":0.0051893042,"about_ca_topic_score_gemma":0.008888734,"teacher_disagreement_score":0.0051893042,"about_ca_system_score_codex":0.00096262293,"about_ca_system_score_gemma":0.0020057543,"threshold_uncertainty_score":0.014716327},"labels":[],"label_agreement":null},{"id":"W4313578599","doi":"10.1002/smr.2529","title":"Optimized fuzzy clustering‐based k‐nearest neighbors imputation for mixed missing data in software development effort estimation","year":2023,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Data mining; Computer science; Categorical variable; Imputation (statistics); Missing data; Cluster analysis; Fuzzy logic; Software; Artificial intelligence; Machine learning","score_opus":0.03699946842914301,"score_gpt":0.3226420344798559,"score_spread":0.28564256605071286,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313578599","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12617528,0.00058269873,0.87101066,0.00023551534,0.00005483575,0.00010229826,0.00032210536,0.00057629304,0.0009402742],"genre_scores_gemma":[0.6675017,0.00017833785,0.33062798,0.00007456196,0.000029778843,0.0001065541,0.0007188871,0.000054336208,0.0007078679],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99444854,0.0028928316,0.0003688221,0.000944176,0.0010884254,0.00025723255],"domain_scores_gemma":[0.98299754,0.010657638,0.0015360147,0.0014458533,0.0030956706,0.0002672069],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010425247,0.0007243483,0.0015866683,0.0022659989,0.00081195415,0.0012782607,0.0022021981,0.0011617928,0.0008964561],"category_scores_gemma":[0.028262569,0.00049970817,0.0014238023,0.0021884576,0.0005022414,0.0014887382,0.0010408137,0.0012952535,0.00032436382],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045931237,0.00025406096,0.01831599,0.00025191705,0.0004112027,0.00013594623,0.000335293,0.7963141,0.0010223154,0.003293098,0.001674696,0.17753208],"study_design_scores_gemma":[0.000012123963,0.000039730283,0.0022866612,0.000027716267,0.000030911415,0.000025536125,0.00006559455,0.99436945,0.0008735862,0.0020005016,0.0002509391,0.000017201119],"about_ca_topic_score_codex":0.015475804,"about_ca_topic_score_gemma":0.012478202,"teacher_disagreement_score":0.015475804,"about_ca_system_score_codex":0.001260939,"about_ca_system_score_gemma":0.001969549,"threshold_uncertainty_score":0.055134654},"labels":[],"label_agreement":null},{"id":"W4313648295","doi":"10.21203/rs.3.rs-2434133/v1","title":"Adaptive Fine-tuning for Multiclass Classification over Software Requirement Data","year":2023,"lang":"en","type":"preprint","venue":"Research Square","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Transfer of learning; Software; Sentence; Pooling; Machine learning; Field (mathematics); Robustness (evolution); Embedding; Deep learning; Task (project management); Natural language processing","score_opus":0.4032188134488606,"score_gpt":0.46393716617561753,"score_spread":0.06071835272675691,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313648295","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12508808,0.0017035402,0.86207867,0.000724046,0.00025607835,0.0001941357,0.00068528287,0.007994615,0.0012756091],"genre_scores_gemma":[0.72067785,0.00024992516,0.27300575,0.0006544962,0.00024061475,0.00022978614,0.0018106479,0.00049755484,0.0026332876],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99703985,0.0008673338,0.0002660946,0.0010528573,0.0004137449,0.00036011374],"domain_scores_gemma":[0.990729,0.005705289,0.0004895713,0.0018461713,0.000948504,0.00028160162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00505684,0.0009570409,0.0021803624,0.0016668235,0.0007463639,0.0012229735,0.0022860288,0.001870652,0.002039045],"category_scores_gemma":[0.015255089,0.00056147046,0.0012507424,0.0017792205,0.0006158524,0.0015681552,0.0016361775,0.00297557,0.0010486286],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078200584,0.0008647488,0.009034363,0.00016665694,0.00024060246,0.000079259225,0.0001590176,0.13737965,0.01611883,0.0011334087,0.010282646,0.8237587],"study_design_scores_gemma":[0.0000374055,0.000061334336,0.002143276,0.000013996475,0.000028199147,0.000039716648,0.00003901214,0.99123,0.0022412715,0.0035895652,0.0005640284,0.000012267394],"about_ca_topic_score_codex":0.009896573,"about_ca_topic_score_gemma":0.01654852,"teacher_disagreement_score":0.009896573,"about_ca_system_score_codex":0.0011204957,"about_ca_system_score_gemma":0.0015647992,"threshold_uncertainty_score":0.026743472},"labels":[],"label_agreement":null},{"id":"W4315815576","doi":"10.1109/scam55253.2022.00012","title":"Revisiting the Impact of Anti-patterns on Fault-Proneness: A Differentiated Replication","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Heuristics; Code refactoring; Computer science; Software quality; Replication (statistics); Code smell; Leverage (statistics); Empirical research; Software; Data mining; Software fault tolerance; Machine learning; Artificial intelligence; Software development; Programming language; Statistics; Mathematics","score_opus":0.029333953436030704,"score_gpt":0.3143885704610985,"score_spread":0.28505461702506785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4315815576","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9222547,0.014203293,0.0451885,0.0033094217,0.0007178208,0.0010547983,0.0070265452,0.0012767881,0.0049681845],"genre_scores_gemma":[0.9592973,0.0009920065,0.03057461,0.0014554334,0.00027937818,0.00044990453,0.0053651817,0.00043636237,0.0011497022],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9676326,0.015329621,0.0028822906,0.008451745,0.0051357206,0.0005680235],"domain_scores_gemma":[0.57382995,0.2504195,0.022668988,0.11096794,0.03976234,0.0023512004],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.036481712,0.0015215346,0.0015163039,0.004361583,0.001198959,0.0040818,0.003881127,0.002134691,0.003781207],"category_scores_gemma":[0.21227898,0.0009448059,0.0031856315,0.004036331,0.0023636837,0.0057555297,0.003392576,0.0027858552,0.0020081955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003038144,0.0022093402,0.72594047,0.005586694,0.0053892643,0.001357276,0.008803461,0.009888303,0.023115551,0.003468786,0.015965616,0.19523712],"study_design_scores_gemma":[0.0011836418,0.006434533,0.83073497,0.0030220395,0.006583365,0.0020292662,0.0105299335,0.06848722,0.02166825,0.011816185,0.036908813,0.00060183223],"about_ca_topic_score_codex":0.014471894,"about_ca_topic_score_gemma":0.011174422,"teacher_disagreement_score":0.96351826,"about_ca_system_score_codex":0.0011342062,"about_ca_system_score_gemma":0.0021820487,"threshold_uncertainty_score":0.192936},"labels":[],"label_agreement":null},{"id":"W4315926717","doi":"10.5281/zenodo.7533156","title":"A Multi-Step Learning Approach to Assist Code Review","year":2022,"lang":"en","type":"paratext","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Code (set theory); Programming language; Code review; Software engineering; Static program analysis; Software; Software development","score_opus":0.059864396724408686,"score_gpt":0.28508928827031205,"score_spread":0.22522489154590336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4315926717","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012568861,0.0001630061,0.93453306,0.0025924847,0.00031073182,0.0029189144,0.0005227432,0.02070784,0.025682395],"genre_scores_gemma":[0.040945582,0.00012211963,0.93243515,0.0005429457,0.00007803449,0.001593487,0.0008628459,0.00068694755,0.022732887],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.986047,0.0074141463,0.00074265763,0.0015605575,0.0037909742,0.00044465472],"domain_scores_gemma":[0.9492957,0.028866274,0.0020314998,0.0042866687,0.012791247,0.0027285272],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010061706,0.0017205185,0.00080279313,0.0034915975,0.0017699113,0.0036816855,0.0044925367,0.0029325094,0.03353245],"category_scores_gemma":[0.046465654,0.00079417543,0.0008055903,0.0018799199,0.0008330696,0.004676497,0.008150658,0.0023716695,0.013585996],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034955668,0.0013825517,0.0011568547,0.0005126953,0.000038938357,0.00030677154,0.003974357,0.004681946,0.007558904,0.005323656,0.07137234,0.9033415],"study_design_scores_gemma":[0.00095558795,0.0027520184,0.006945056,0.0016609785,0.00019442795,0.0026378967,0.0063812183,0.2657882,0.058477912,0.08395632,0.5696668,0.0005835637],"about_ca_topic_score_codex":0.0014484603,"about_ca_topic_score_gemma":0.0054357555,"teacher_disagreement_score":0.03353245,"about_ca_system_score_codex":0.001583596,"about_ca_system_score_gemma":0.0047080754,"threshold_uncertainty_score":0.11217719},"labels":[],"label_agreement":null},{"id":"W4317209725","doi":"10.1145/3576037","title":"I Depended on You and You Broke Me: An Empirical Study of Manifesting Breaking Changes in Client Packages","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Software versioning; Breaking strength; Change impact analysis; Dependency (UML); Software; Software engineering; Operating system","score_opus":0.14032214583660246,"score_gpt":0.3833204554910334,"score_spread":0.24299830965443092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4317209725","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9977622,0.00009842802,0.0007596356,0.00020927726,0.0000052673686,0.0000385314,0.000071173126,0.00003561544,0.0010199656],"genre_scores_gemma":[0.9982059,0.00007836037,0.0009085274,0.00013329294,0.000008717967,0.00003930329,0.00016086541,0.000051636012,0.000413389],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9838095,0.008996811,0.0012568786,0.0015312281,0.0035316,0.00087397435],"domain_scores_gemma":[0.71232516,0.2042191,0.04604715,0.018457396,0.014126418,0.0048248502],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016221542,0.0004553901,0.00040603275,0.0017882022,0.0016456963,0.0025197829,0.0017825171,0.001538369,0.0021473067],"category_scores_gemma":[0.1562494,0.00074580545,0.0003840088,0.0020672497,0.0031198661,0.006903557,0.00275074,0.0036624137,0.0007701556],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042777666,0.0015571672,0.86349803,0.00026240695,0.000118349424,0.0013487634,0.089617066,0.0006134678,0.0020487986,0.0012983491,0.0019867886,0.037223063],"study_design_scores_gemma":[0.000041837873,0.00071066455,0.9283539,0.00020657231,0.00007734775,0.001557556,0.053700846,0.0062521813,0.0014211325,0.0014087552,0.0061733094,0.000095933305],"about_ca_topic_score_codex":0.0057940655,"about_ca_topic_score_gemma":0.0059509533,"teacher_disagreement_score":0.016221542,"about_ca_system_score_codex":0.0014964467,"about_ca_system_score_gemma":0.001192208,"threshold_uncertainty_score":0.08578873},"labels":[],"label_agreement":null},{"id":"W4318148225","doi":"10.1109/bigdata55660.2022.10020386","title":"Adaptive Method for Machine Learning Model Selection in Data Science Projects","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Big Data (Big Data)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Selection (genetic algorithm); Machine learning; Process (computing); Heuristics; Feature selection; Model selection; Artificial intelligence; Data mining","score_opus":0.4726800557815008,"score_gpt":0.4217361885031067,"score_spread":0.05094386727839412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4318148225","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0052037993,0.000044003194,0.99351317,0.00010168571,0.000013121596,0.00010563633,0.000041332525,0.0005492142,0.000428109],"genre_scores_gemma":[0.1561818,0.000062827356,0.8414322,0.0001569201,0.000033153538,0.000732776,0.00029894692,0.00020844261,0.00089298707],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9905923,0.0060970583,0.0004585738,0.001169952,0.0014121787,0.00026998724],"domain_scores_gemma":[0.9591773,0.03396612,0.0015706968,0.0022822358,0.0026006137,0.0004029306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01356869,0.0010950305,0.00091748656,0.0029287112,0.0009752389,0.0018309184,0.0026372445,0.0016391269,0.0038153096],"category_scores_gemma":[0.048986614,0.0007623716,0.0017023725,0.0022778574,0.0009944342,0.0021379853,0.0023432525,0.0033485962,0.00066401967],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041866204,0.00037942908,0.012005607,0.00032979556,0.0004161185,0.00041006313,0.00082984666,0.5824137,0.0031299528,0.06175525,0.004448804,0.33346286],"study_design_scores_gemma":[0.000026063632,0.000030244668,0.00026473304,0.000016084405,0.00001652133,0.00002964246,0.00004130354,0.983095,0.00062661717,0.014736877,0.0011054692,0.000011431773],"about_ca_topic_score_codex":0.004615558,"about_ca_topic_score_gemma":0.0074260947,"teacher_disagreement_score":0.01356869,"about_ca_system_score_codex":0.0015683426,"about_ca_system_score_gemma":0.0028612197,"threshold_uncertainty_score":0.071758926},"labels":[],"label_agreement":null},{"id":"W4319083508","doi":"10.1007/s10664-022-10276-6","title":"An empirical study of text-based machine learning models for vulnerability detection","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Machine learning; Vulnerability (computing); Artificial intelligence; Context (archaeology); Empirical research; Source code; Function (biology); Construct (python library); Vulnerability assessment; Code (set theory); Data science; Computer security; Geography; Programming language","score_opus":0.059584817812828254,"score_gpt":0.3385498813244283,"score_spread":0.27896506351160005,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319083508","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97985816,0.0007246525,0.015138775,0.000889964,0.000078129735,0.000096510834,0.0008810761,0.00017491245,0.0021577803],"genre_scores_gemma":[0.99229395,0.0001983368,0.005103146,0.00013869347,0.00008455763,0.000060759718,0.001279091,0.000051042767,0.0007904144],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.992922,0.0049486435,0.00043053832,0.00067849684,0.000851566,0.00016883937],"domain_scores_gemma":[0.5946221,0.38334325,0.0081098545,0.006657025,0.0064205215,0.00084743445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0126240095,0.00085677,0.0005776352,0.002203456,0.00057917845,0.0018476757,0.0015777835,0.0017187053,0.0030322343],"category_scores_gemma":[0.14984278,0.0002856331,0.0006630123,0.0029152827,0.0007572208,0.005758723,0.0009314216,0.0023911395,0.0011882053],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004864189,0.008742703,0.52722716,0.0010906134,0.0009737028,0.00089739484,0.0032689979,0.11559635,0.004655746,0.008639983,0.013941054,0.31010222],"study_design_scores_gemma":[0.0001544081,0.0010560628,0.084546514,0.00015208121,0.00029682097,0.0005363505,0.0009737974,0.89776814,0.0026227369,0.009314727,0.002503626,0.000074750875],"about_ca_topic_score_codex":0.0040455637,"about_ca_topic_score_gemma":0.0029542572,"teacher_disagreement_score":0.0126240095,"about_ca_system_score_codex":0.0011769874,"about_ca_system_score_gemma":0.0006639501,"threshold_uncertainty_score":0.066762924},"labels":[],"label_agreement":null},{"id":"W4319264802","doi":"10.1016/j.infsof.2023.107169","title":"Just-in-time code duplicates extraction","year":2023,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; Plug-in; Workflow; Program slicing; Artificial intelligence; Source code; Machine learning; Software engineering; Deep learning; Coding (social sciences); Code review; Classifier (UML); Software maintenance; Convolutional neural network; Software; Static program analysis; Programming language; Software development; Database","score_opus":0.013380135617879992,"score_gpt":0.27401743899741166,"score_spread":0.26063730337953167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319264802","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04723328,0.0016587533,0.8835029,0.00092863146,0.0011360223,0.0006834851,0.0061861197,0.04872809,0.009942824],"genre_scores_gemma":[0.19721596,0.0008712372,0.7540732,0.0004917858,0.00034280657,0.00032643828,0.016585657,0.0069008237,0.023192149],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99416566,0.00060748216,0.0005611737,0.0010118423,0.0031262792,0.00052759383],"domain_scores_gemma":[0.9855816,0.003045892,0.0008556741,0.005763393,0.004413008,0.00034038877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013537813,0.0019542435,0.0017907378,0.0046779397,0.0017091094,0.0032752545,0.0029639108,0.0016298014,0.009482614],"category_scores_gemma":[0.016645938,0.0009335939,0.002247276,0.0036083318,0.0008263506,0.004475105,0.004131313,0.0013973382,0.009179757],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007343833,0.00017094011,0.0044811433,0.0015252444,0.00021581606,0.0011537144,0.00052621006,0.0044886773,0.054080747,0.013959818,0.060353473,0.85830986],"study_design_scores_gemma":[0.0003025949,0.0006429669,0.008483179,0.00046933882,0.00067782175,0.0065914043,0.0013116217,0.2553264,0.3761458,0.08819899,0.26148635,0.0003636401],"about_ca_topic_score_codex":0.0029109602,"about_ca_topic_score_gemma":0.0059744627,"teacher_disagreement_score":0.009482614,"about_ca_system_score_codex":0.0007967307,"about_ca_system_score_gemma":0.0051805703,"threshold_uncertainty_score":0.031722546},"labels":[],"label_agreement":null},{"id":"W4319451811","doi":"10.48550/arxiv.2206.05480","title":"CodeS: Towards Code Model Generalization Under Distribution Shift","year":2022,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"JST-Mirai Program; Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Python (programming language); Abstract syntax tree; Source code; Programming language; Code refactoring; Syntax; Code (set theory); Java; Artificial intelligence; Software","score_opus":0.09117559834616137,"score_gpt":0.22586006778830547,"score_spread":0.1346844694421441,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319451811","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1791732,0.0014855758,0.7825469,0.0029323963,0.00033757192,0.00025081288,0.0041473797,0.024406651,0.0047195703],"genre_scores_gemma":[0.76748765,0.00070058863,0.19976339,0.0021376642,0.00025237745,0.00048658805,0.020406973,0.0025303334,0.0062344545],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99639434,0.0008854599,0.00017479408,0.0014951251,0.0007464018,0.00030381684],"domain_scores_gemma":[0.9882151,0.004838178,0.0007373759,0.004083439,0.0017234124,0.00040257323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004707425,0.0020418833,0.00127083,0.0020440472,0.00087003125,0.0023653496,0.0031639661,0.0025139758,0.0019405484],"category_scores_gemma":[0.028487045,0.0007516655,0.0018365915,0.0014753944,0.0018298589,0.005027157,0.0043845903,0.0061137807,0.0021841158],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076362246,0.00042136988,0.020091575,0.00043170867,0.00031382445,0.00029636573,0.00047502655,0.4563662,0.012166669,0.021198273,0.048225325,0.43924996],"study_design_scores_gemma":[0.000031367945,0.000073030526,0.0007567271,0.000025042049,0.000020323097,0.000069686146,0.0000553612,0.96347743,0.0040926603,0.028472222,0.002910744,0.000015378586],"about_ca_topic_score_codex":0.010034538,"about_ca_topic_score_gemma":0.009480737,"teacher_disagreement_score":0.010034538,"about_ca_system_score_codex":0.0025649276,"about_ca_system_score_gemma":0.0032412892,"threshold_uncertainty_score":0.024895549},"labels":[],"label_agreement":null},{"id":"W4319459392","doi":"10.1007/s10664-022-10270-y","title":"Assessing the exposure of software changes","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Deliverable; Computer science; Ranking (information retrieval); Software; Source code; Task (project management); Code (set theory); Set (abstract data type); Software engineering; Reliability engineering; Data mining; Systems engineering; Engineering; Machine learning; Operating system; Programming language","score_opus":0.050832639175595495,"score_gpt":0.3292309896640568,"score_spread":0.2783983504884613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319459392","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99514365,0.00008467989,0.0014829638,0.00010821069,0.0000067990595,0.000029417057,0.000116245676,0.000026138683,0.0030019365],"genre_scores_gemma":[0.99840564,0.000059446735,0.0007409737,0.000026817801,0.00000926256,0.000015792066,0.00014965718,0.0000061334677,0.00058624893],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99299616,0.0024018646,0.00060785055,0.0006799315,0.0029218148,0.0003923766],"domain_scores_gemma":[0.86608607,0.08814725,0.027384745,0.0067336857,0.008198123,0.0034501904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052902596,0.00028921932,0.00025944403,0.0027849982,0.0003688872,0.0013126407,0.0006488798,0.0012004274,0.004065148],"category_scores_gemma":[0.098083735,0.00020502097,0.0004941428,0.001461487,0.0004986171,0.002455294,0.0014711707,0.001079174,0.0005355699],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037378422,0.00089549785,0.93878716,0.000071512724,0.00017543165,0.000082653125,0.0011503345,0.001519352,0.0017695104,0.00078155217,0.00023949066,0.054153647],"study_design_scores_gemma":[0.000020144944,0.0016212511,0.9836838,0.000038703143,0.000093404225,0.0002589593,0.001727269,0.0069231945,0.002638925,0.0017559201,0.0012121939,0.000026287194],"about_ca_topic_score_codex":0.0014085177,"about_ca_topic_score_gemma":0.0016092929,"teacher_disagreement_score":0.0052902596,"about_ca_system_score_codex":0.0006388686,"about_ca_system_score_gemma":0.00058374647,"threshold_uncertainty_score":0.027977884},"labels":[],"label_agreement":null},{"id":"W4320882938","doi":"10.1007/s10664-022-10271-x","title":"Refactoring practices in the context of data-intensive systems","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Data access; Context (archaeology); Maintainability; Database; Software engineering; Software; Programming language","score_opus":0.16027050351680053,"score_gpt":0.3800548519220058,"score_spread":0.21978434840520525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4320882938","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9960426,0.00058974425,0.0011032347,0.0004754391,0.000006115136,0.000019485937,0.00003445896,0.000015365506,0.0017134441],"genre_scores_gemma":[0.9980469,0.00025481687,0.0014369104,0.00003426858,0.0000037302905,0.000005742885,0.000031834057,0.000007524629,0.00017830936],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99152595,0.0050004474,0.00061073346,0.0008438673,0.00149968,0.0005192598],"domain_scores_gemma":[0.88862044,0.08249551,0.0114788655,0.00533265,0.0096119605,0.0024606124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008906294,0.0002451576,0.00018764618,0.0021357932,0.0015275613,0.0026189452,0.0010437464,0.0010454267,0.0013757874],"category_scores_gemma":[0.07565498,0.00028182066,0.0001742249,0.002892689,0.0010288983,0.0032463477,0.0016845809,0.0011693136,0.00017381342],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007624809,0.002025166,0.5039103,0.0009570632,0.000170851,0.0026367416,0.13914487,0.008526495,0.014254497,0.019302396,0.0019310232,0.30637816],"study_design_scores_gemma":[0.00014288169,0.0019853632,0.7322642,0.0015861065,0.00028675448,0.002654128,0.15440968,0.03272543,0.015359051,0.027242359,0.03113699,0.00020704539],"about_ca_topic_score_codex":0.010212273,"about_ca_topic_score_gemma":0.027542293,"teacher_disagreement_score":0.010212273,"about_ca_system_score_codex":0.0029193899,"about_ca_system_score_gemma":0.003059176,"threshold_uncertainty_score":0.047101498},"labels":[],"label_agreement":null},{"id":"W4321022177","doi":"10.2139/ssrn.4361607","title":"Auditing Large Language Models: A Three-Layered Approach","year":2023,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute on Governance","funders":"","keywords":"Audit; Corporate governance; Blueprint; Engineering ethics; Business; Engineering; Accounting; Finance","score_opus":0.01955062911801558,"score_gpt":0.2626686702789555,"score_spread":0.24311804116093993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321022177","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016968606,0.00020765084,0.97558933,0.000864027,0.000043736527,0.00030391148,0.00014204462,0.003731214,0.002149463],"genre_scores_gemma":[0.3384138,0.00030007775,0.6556831,0.00031140097,0.000094028364,0.00024924387,0.00057793694,0.000871652,0.0034987573],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.95826745,0.0132576395,0.004430926,0.0028006062,0.018248469,0.0029948996],"domain_scores_gemma":[0.8725809,0.040621288,0.007950001,0.057308517,0.018804824,0.0027344802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02683224,0.0013448655,0.002488328,0.0052903197,0.0025812394,0.013954208,0.006731809,0.0034421564,0.00397595],"category_scores_gemma":[0.09875196,0.0028150787,0.0034669985,0.0034700013,0.003518296,0.024713803,0.017813431,0.0057043764,0.0013669541],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013707405,0.0009264445,0.027951393,0.00074321026,0.0008356987,0.0011190688,0.0037224186,0.12969147,0.023824131,0.18744873,0.008218477,0.6141482],"study_design_scores_gemma":[0.000049896826,0.00016676192,0.0016528044,0.00017494174,0.00024995118,0.00040642646,0.000642406,0.8046535,0.015374953,0.17124599,0.005221384,0.00016095178],"about_ca_topic_score_codex":0.008264806,"about_ca_topic_score_gemma":0.013198324,"teacher_disagreement_score":0.02683224,"about_ca_system_score_codex":0.0024439632,"about_ca_system_score_gemma":0.007712827,"threshold_uncertainty_score":0.14190412},"labels":[],"label_agreement":null},{"id":"W4321061955","doi":"10.1109/apsec57359.2022.00088","title":"A new measure to assess the systematicity of the abstracts of reviews self-identifying as systematic reviews","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Systematic review; Computer science; Quality (philosophy); Field (mathematics); Data science; Measure (data warehouse); Management science; Engineering ethics; MEDLINE; Epistemology; Political science; Data mining; Engineering","score_opus":0.13287832130916608,"score_gpt":0.34009220685148117,"score_spread":0.2072138855423151,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321061955","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05793934,0.17760764,0.6620949,0.010055556,0.0045712455,0.03164142,0.032333314,0.0034488346,0.020307707],"genre_scores_gemma":[0.37161022,0.018917473,0.55187947,0.0027877598,0.0023259998,0.041375622,0.009280882,0.0004667604,0.0013557639],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.43638197,0.25956157,0.19282565,0.01983451,0.08956627,0.0018300745],"domain_scores_gemma":[0.10491796,0.7267021,0.08577683,0.028309148,0.05269456,0.0015993344],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.35105437,0.0040691155,0.0071934746,0.06590802,0.0024844501,0.008373046,0.0038219339,0.0061195698,0.005388203],"category_scores_gemma":[0.77172506,0.0019012314,0.015635772,0.045310266,0.006279535,0.017227286,0.008098298,0.003867616,0.0013712347],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0063574333,0.00063158927,0.13209063,0.18778385,0.058576554,0.00093176565,0.01108239,0.008777881,0.0067070434,0.036764797,0.02437922,0.5259168],"study_design_scores_gemma":[0.009464694,0.009754476,0.24531032,0.070669584,0.10998204,0.008767988,0.006622246,0.06968152,0.01677731,0.27699852,0.17254275,0.0034285756],"about_ca_topic_score_codex":0.001602265,"about_ca_topic_score_gemma":0.0026393975,"teacher_disagreement_score":0.6489456,"about_ca_system_score_codex":0.006510852,"about_ca_system_score_gemma":0.0109891975,"threshold_uncertainty_score":0.8002655},"labels":[],"label_agreement":null},{"id":"W4321061983","doi":"10.1109/apsec57359.2022.00071","title":"A checklist-based approach to assess the systematicity of the abstracts of reviews self-identifying as systematic reviews","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Systematic review; Checklist; Computer science; Quality (philosophy); Field (mathematics); Management science; Data science; Engineering ethics; Psychology; MEDLINE; Political science; Epistemology; Engineering; Cognitive psychology","score_opus":0.11563327050482876,"score_gpt":0.3291899797183531,"score_spread":0.21355670921352438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321061983","genre_codex":"protocol","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007397077,0.024185136,0.4557863,0.009832933,0.004002624,0.46570686,0.01790001,0.004642395,0.010546686],"genre_scores_gemma":[0.013184358,0.0034668094,0.7520917,0.0011710669,0.0002759358,0.22613749,0.00275358,0.00019603332,0.0007229625],"study_design_codex":"systematic_review","study_design_gemma":"observational","domain_scores_codex":[0.23992078,0.42812625,0.2622209,0.011013997,0.05657982,0.0021382682],"domain_scores_gemma":[0.14690709,0.5556311,0.083819725,0.03929768,0.16945164,0.0048928484],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.5151922,0.008586025,0.012674272,0.07858091,0.0076907985,0.012171884,0.009687406,0.008403773,0.014050225],"category_scores_gemma":[0.72153527,0.0050821076,0.0246509,0.04892991,0.009003795,0.012764661,0.014480262,0.010296234,0.005175648],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033547857,0.0005996617,0.007360257,0.44041327,0.015589602,0.0008218626,0.019132046,0.0023908662,0.0077245184,0.031495646,0.059066687,0.41205084],"study_design_scores_gemma":[0.014045637,0.0051243603,0.02877285,0.27559316,0.041733082,0.0029737393,0.012629458,0.017584544,0.014808917,0.12565342,0.45725867,0.0038221388],"about_ca_topic_score_codex":0.004465867,"about_ca_topic_score_gemma":0.010919696,"teacher_disagreement_score":0.4848078,"about_ca_system_score_codex":0.016634539,"about_ca_system_score_gemma":0.08339153,"threshold_uncertainty_score":0.5978544},"labels":[],"label_agreement":null},{"id":"W4321242415","doi":"10.1007/s10009-023-00696-0","title":"Algorithm selection for SMT","year":2023,"lang":"en","type":"article","venue":"International Journal on Software Tools for Technology Transfer","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Satisfiability modulo theories; Solver; Ranking (information retrieval); Pairwise comparison; Theory of computation; Selection (genetic algorithm); Domain (mathematical analysis); Theoretical computer science; Programming language; Algorithm; Artificial intelligence; Mathematics","score_opus":0.02944465937477168,"score_gpt":0.31299095544415556,"score_spread":0.28354629606938386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321242415","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013271586,0.00043503175,0.97966355,0.0001980821,0.00012242545,0.000117546595,0.0001803527,0.0025338784,0.003477603],"genre_scores_gemma":[0.20496413,0.00024442645,0.7855612,0.00020108807,0.00015992543,0.00038711133,0.0012766337,0.0007850222,0.006420453],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99847144,0.00060355803,0.00011516536,0.00026638166,0.00039479346,0.00014855202],"domain_scores_gemma":[0.9970921,0.0017978885,0.00011315181,0.00032534945,0.00057858485,0.0000928962],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017775029,0.0010461779,0.0012439046,0.0020495825,0.0010084915,0.0012612178,0.0013612313,0.0016616494,0.011422779],"category_scores_gemma":[0.007770849,0.00049127726,0.001438336,0.0014618357,0.00042184725,0.0012505489,0.0012281612,0.0013229569,0.0033359255],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047182618,0.00015949157,0.0013979,0.00021222298,0.00016247325,0.00020121592,0.00008259697,0.17677274,0.0071162377,0.010865506,0.011574546,0.7909832],"study_design_scores_gemma":[0.000053675463,0.0001131464,0.00031519154,0.000019992982,0.000029429908,0.00011699851,0.00002264778,0.9840189,0.0028054577,0.008907444,0.003588189,0.000008892874],"about_ca_topic_score_codex":0.0018123679,"about_ca_topic_score_gemma":0.0026680976,"teacher_disagreement_score":0.011422779,"about_ca_system_score_codex":0.0005714404,"about_ca_system_score_gemma":0.001640369,"threshold_uncertainty_score":0.038213015},"labels":[],"label_agreement":null},{"id":"W4321373549","doi":"10.1017/s0890060422000269","title":"Graph models for engineering design: Model encoding, and fidelity evaluation based on dataset and other sources of knowledge","year":2023,"lang":"en","type":"article","venue":"Artificial intelligence for engineering design analysis and manufacturing","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"University of Cambridge","keywords":"Computer science; Fidelity; Machine learning; Graph; Artificial intelligence; Data modeling; Data mining; Knowledge extraction; Theoretical computer science; Software engineering","score_opus":0.1348726931625898,"score_gpt":0.33442678356376837,"score_spread":0.19955409040117858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321373549","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08297598,0.00046891175,0.9101604,0.0008252375,0.00003222623,0.00028210846,0.0017350547,0.0012150838,0.0023050432],"genre_scores_gemma":[0.61681724,0.0003519702,0.37866503,0.00011086403,0.000022368627,0.00041834204,0.0030763424,0.00011358058,0.00042421024],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99528015,0.0030048518,0.00029346233,0.00049144286,0.0008384771,0.0000916414],"domain_scores_gemma":[0.9473926,0.04305194,0.0022824581,0.0048656347,0.002131447,0.00027588668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006911208,0.000666798,0.00064157846,0.0040326463,0.0004294864,0.0024152396,0.0011209517,0.0012560287,0.0017992262],"category_scores_gemma":[0.054559134,0.0004068352,0.0012976187,0.00270866,0.001004715,0.0033090536,0.0016336598,0.0013737169,0.00022094444],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021767725,0.00021331395,0.008102777,0.00033670204,0.00012873135,0.00010421107,0.00029750768,0.8398789,0.0012492223,0.038258366,0.0013372183,0.10987542],"study_design_scores_gemma":[0.000010364286,0.00003282866,0.0007097834,0.00004360653,0.000019568632,0.00002015789,0.00003927278,0.9772298,0.0006045291,0.020471305,0.0008093768,0.000009358942],"about_ca_topic_score_codex":0.0064785634,"about_ca_topic_score_gemma":0.0073426706,"teacher_disagreement_score":0.006911208,"about_ca_system_score_codex":0.00213151,"about_ca_system_score_gemma":0.001359722,"threshold_uncertainty_score":0.036550403},"labels":[],"label_agreement":null},{"id":"W4321609628","doi":"10.1109/jsyst.2023.3241415","title":"Editorial","year":2023,"lang":"es","type":"editorial","venue":"IEEE Systems Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Code (set theory); Machine learning; Artificial intelligence; Training set; Significant difference; Data mining; Statistics; Mathematics","score_opus":0.020986463531106343,"score_gpt":0.2971159760996975,"score_spread":0.27612951256859114,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321609628","genre_codex":"other","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013110685,0.014003028,0.005597716,0.061221886,0.2509622,0.00033737338,0.0049482207,0.003072681,0.65854585],"genre_scores_gemma":[0.012909639,0.011638864,0.0034643547,0.027548688,0.073601305,0.00022982224,0.0054691164,0.0018205665,0.86331767],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99626666,0.000559455,0.00029651826,0.00092081475,0.0015904133,0.00036617945],"domain_scores_gemma":[0.98957497,0.0019105147,0.0006090042,0.0015952006,0.0042998884,0.002010382],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0025776632,0.001206611,0.0010552382,0.0026761694,0.0023897465,0.009015502,0.002966639,0.0039103357,0.51215965],"category_scores_gemma":[0.018251762,0.000494446,0.0012581204,0.0017191768,0.0013847788,0.00578333,0.0037702625,0.0041712383,0.3246138],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044051383,0.000022519338,0.00020768959,0.00029186206,0.0000095087535,0.00015104047,0.00009227225,0.0000685226,0.00018370575,0.007783993,0.9117943,0.07935052],"study_design_scores_gemma":[0.0000050025087,0.0000064636415,0.000098754135,0.000110303285,0.0000032306111,0.00010208414,0.00005013353,0.000032098436,0.00007226826,0.0011617442,0.9983537,0.0000042480206],"about_ca_topic_score_codex":0.0020804289,"about_ca_topic_score_gemma":0.002717318,"teacher_disagreement_score":0.48784035,"about_ca_system_score_codex":0.0021305294,"about_ca_system_score_gemma":0.0042185416,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4321787049","doi":"10.1007/s10664-022-10272-w","title":"Semantically-enhanced topic recommendation systems for software projects","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Metadata; Recommender system; Software; Information retrieval; World Wide Web; Software engineering; Data science; Programming language","score_opus":0.050022828276608984,"score_gpt":0.3143586557869092,"score_spread":0.2643358275103002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321787049","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06399181,0.0018029597,0.8800558,0.0007802607,0.000418434,0.00046227706,0.0052124322,0.04284619,0.0044299015],"genre_scores_gemma":[0.2602555,0.0008125278,0.7190115,0.0002598594,0.00034441391,0.00030485925,0.012942291,0.0012438429,0.0048251706],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99725133,0.0008534176,0.00031320777,0.00064758153,0.00072767097,0.00020677855],"domain_scores_gemma":[0.9932594,0.0026879439,0.00039303498,0.001713409,0.0015878117,0.0003584253],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029212378,0.00097093027,0.0011704392,0.004982917,0.0010093796,0.0022736238,0.0014993658,0.0015157956,0.005274057],"category_scores_gemma":[0.015026502,0.00048572963,0.0015780855,0.0038787888,0.00034476363,0.0041179606,0.0030700294,0.0015259599,0.0042357426],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016184907,0.00070169894,0.012599644,0.0008899441,0.0004949311,0.00030651301,0.0012202861,0.01792742,0.027793638,0.013790015,0.04428585,0.87837154],"study_design_scores_gemma":[0.00047684423,0.00066151953,0.010202222,0.00023498763,0.00069301337,0.00071296276,0.0008571367,0.8372187,0.02743153,0.06222538,0.05906446,0.0002212998],"about_ca_topic_score_codex":0.0051426496,"about_ca_topic_score_gemma":0.013287222,"teacher_disagreement_score":0.005274057,"about_ca_system_score_codex":0.00078458217,"about_ca_system_score_gemma":0.0015026518,"threshold_uncertainty_score":0.017643452},"labels":[],"label_agreement":null},{"id":"W4321793856","doi":"10.1016/j.jss.2023.111651","title":"Finding associations between natural and computer languages: A case-study of bilingual LDA applied to the bleeping computer forum posts","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada); Lakehead University; Queen's University","funders":"","keywords":"Perplexity; Computer science; Topic model; Natural language processing; Context (archaeology); Latent Dirichlet allocation; Artificial intelligence; Coherence (philosophical gambling strategy); Natural language; Software; Language model; Statistics; Mathematics; Programming language","score_opus":0.026798585533271354,"score_gpt":0.2989785756976393,"score_spread":0.27217999016436795,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321793856","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9914596,0.00015337626,0.0046492973,0.0002340208,0.000012223808,0.00006819245,0.00024640065,0.000055201133,0.0031216287],"genre_scores_gemma":[0.99389416,0.00008414612,0.0043180776,0.000054254368,0.000009782821,0.000041846473,0.00029826871,0.000036110814,0.0012632197],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99674547,0.0021165218,0.0001464777,0.0003513901,0.00038343226,0.00025661927],"domain_scores_gemma":[0.97593415,0.019495493,0.0010035174,0.0009818849,0.0018743445,0.0007104961],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037516118,0.00027920588,0.0003922427,0.0032458634,0.0032372798,0.001554316,0.0005835404,0.00087141595,0.002392235],"category_scores_gemma":[0.020513268,0.00019819704,0.0002445469,0.0035729967,0.0012041137,0.0021730878,0.0018023864,0.000883284,0.00058551953],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018493714,0.0018043837,0.6115392,0.00089761545,0.00012509266,0.0055178576,0.16041718,0.0007783431,0.024620226,0.0051622973,0.0040042177,0.18328425],"study_design_scores_gemma":[0.00017506951,0.0011015617,0.5811534,0.00040560574,0.00037242417,0.010382388,0.29346633,0.030515213,0.024653625,0.010390431,0.04716163,0.00022243173],"about_ca_topic_score_codex":0.017809376,"about_ca_topic_score_gemma":0.043090258,"teacher_disagreement_score":0.017809376,"about_ca_system_score_codex":0.0008790951,"about_ca_system_score_gemma":0.0017984573,"threshold_uncertainty_score":0.035411477},"labels":[],"label_agreement":null},{"id":"W4323042475","doi":"10.1145/3585386","title":"VulANalyzeR: Explainable Binary Vulnerability Detection with Multi-task Learning and Attentional Graph Convolution","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Privacy and Security","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; McGill University; Queen's University","funders":"","keywords":"Computer science; Leverage (statistics); Vulnerability (computing); Machine learning; Artificial intelligence; Binary number; Binary classification; Deep learning; Task (project management); Software; Vulnerability assessment; Workload; Graph; Data mining; Theoretical computer science; Computer security; Support vector machine; Engineering","score_opus":0.020615412110099443,"score_gpt":0.27178697453503053,"score_spread":0.2511715624249311,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323042475","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13620101,0.0016679605,0.82843477,0.0011162875,0.00018100473,0.00018944818,0.001366283,0.028025366,0.0028178128],"genre_scores_gemma":[0.78568345,0.000496489,0.20248592,0.00085510826,0.000099861616,0.00022613365,0.003457724,0.00056086114,0.0061345114],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997037,0.00004811079,0.000010922962,0.00012748582,0.000051001753,0.000058640213],"domain_scores_gemma":[0.9993259,0.0003515103,0.00007650391,0.000114696784,0.000081313745,0.000050093076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006397241,0.0016860603,0.0008433942,0.00097603136,0.00033807594,0.00077910506,0.0022426748,0.001671694,0.0022403128],"category_scores_gemma":[0.0023291293,0.0006179452,0.0011531705,0.0006435476,0.0005643199,0.0016465082,0.001415986,0.0019598578,0.000704738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041944213,0.00042537146,0.005477302,0.00023121333,0.00027769784,0.00034762197,0.00013307433,0.5445051,0.011843399,0.00493089,0.018103605,0.41330525],"study_design_scores_gemma":[0.000010029855,0.000025450387,0.00020777638,0.0000038313137,0.000009952523,0.000018067363,0.000004668695,0.99583983,0.0009989311,0.002566468,0.000309735,0.0000051551333],"about_ca_topic_score_codex":0.012595009,"about_ca_topic_score_gemma":0.01662695,"teacher_disagreement_score":0.012595009,"about_ca_system_score_codex":0.0011785738,"about_ca_system_score_gemma":0.0012094382,"threshold_uncertainty_score":0.025043368},"labels":[],"label_agreement":null},{"id":"W4323526456","doi":"10.1145/3545947.3576298","title":"Impact of Group Member Prerequisite Grades on Problem Set and Test Grades","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Teamwork; Mathematics education; Test (biology); Set (abstract data type); Group work; Work (physics); Group (periodic table); Computer science; Psychology; Medical education; Engineering; Management; Programming language","score_opus":0.026471862582437153,"score_gpt":0.2982271620626502,"score_spread":0.27175529948021304,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323526456","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9978236,0.00003328056,0.00024006951,0.00006461482,0.000018768511,0.000016967906,0.000048078688,0.000031707328,0.0017229812],"genre_scores_gemma":[0.9985297,0.000010138299,0.00016923953,0.00001450788,0.0000048363313,0.000016392865,0.000108135355,0.00000800562,0.0011390733],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99649894,0.0014484661,0.00024039003,0.0005596231,0.0008088481,0.00044366805],"domain_scores_gemma":[0.9433107,0.028122235,0.0053317277,0.0041098837,0.003846945,0.015278432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039763586,0.0005074266,0.0005553021,0.0008231229,0.00068945775,0.0019196386,0.0009274257,0.0006632639,0.008450484],"category_scores_gemma":[0.048037432,0.0001607566,0.00046642747,0.00046282762,0.0005114429,0.000892068,0.0016517319,0.0012509247,0.0016011175],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006795787,0.010037844,0.8829227,0.00006691154,0.00026756965,0.00028585992,0.0011884241,0.0035098873,0.005854503,0.0005097882,0.0018779907,0.086682685],"study_design_scores_gemma":[0.00010693608,0.005220426,0.98796785,0.000020758556,0.00005252784,0.00010666668,0.0007155149,0.0025081355,0.0021356838,0.0004774632,0.0006651477,0.000022808948],"about_ca_topic_score_codex":0.001769256,"about_ca_topic_score_gemma":0.0020610127,"teacher_disagreement_score":0.008450484,"about_ca_system_score_codex":0.00057553884,"about_ca_system_score_gemma":0.00074255455,"threshold_uncertainty_score":0.028269649},"labels":[],"label_agreement":null},{"id":"W4323645873","doi":"10.1109/fnwf55208.2022.00026","title":"A Study on the Use of Runtime Files in Handling Crash Reports in a Large Telecom Company","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ericsson (Canada); Concordia University","funders":"Mitacs","keywords":"Crash; Computer science; Downtime; Service (business); Process (computing); Computer security; Product (mathematics); Operating system; Business","score_opus":0.06143679148979848,"score_gpt":0.28729154688741027,"score_spread":0.2258547553976118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323645873","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99925774,0.000043182048,0.00020539697,0.00008089618,0.0000019018701,0.000025825564,0.000054889475,0.000015830645,0.00031433484],"genre_scores_gemma":[0.9985877,0.00010786777,0.0006384117,0.000064826876,0.0000070377932,0.000021115346,0.0001483281,0.000015390993,0.00040942166],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99608743,0.001870596,0.0003731952,0.0005046736,0.0008077364,0.00035635318],"domain_scores_gemma":[0.8933904,0.07148564,0.013396894,0.0042435084,0.014060136,0.0034233956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051571494,0.00043482674,0.0002532167,0.0026701323,0.0012188932,0.0014868355,0.0010244333,0.0010143883,0.0012927054],"category_scores_gemma":[0.03723772,0.00044587694,0.00027440995,0.0019997242,0.0009062807,0.0018286216,0.00063505745,0.001165852,0.0003928892],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010938897,0.005280731,0.8665947,0.0003383497,0.00020307653,0.0025856958,0.042454977,0.0038308776,0.011740699,0.00049804716,0.0023008122,0.06307815],"study_design_scores_gemma":[0.000055526463,0.003963677,0.9439446,0.00008768323,0.00009668907,0.0011868061,0.03487837,0.008483955,0.0044027884,0.00008379417,0.0027150812,0.000101029],"about_ca_topic_score_codex":0.0230626,"about_ca_topic_score_gemma":0.027213316,"teacher_disagreement_score":0.0230626,"about_ca_system_score_codex":0.002007754,"about_ca_system_score_gemma":0.0012345096,"threshold_uncertainty_score":0.045856714},"labels":[],"label_agreement":null},{"id":"W4323657910","doi":"10.1016/j.infsof.2023.107196","title":"Studying the challenges of developing hardware description language programs","year":2023,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Programming language; Python (programming language); Verilog; Software engineering; Domain (mathematical analysis); World Wide Web; Operating system","score_opus":0.03859914057386985,"score_gpt":0.27184505424048494,"score_spread":0.2332459136666151,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323657910","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18075575,0.0017391385,0.78739065,0.012538563,0.0001670636,0.0004607845,0.00030548355,0.0019154348,0.014727061],"genre_scores_gemma":[0.25195655,0.0014078253,0.74094176,0.0006107227,0.00009354123,0.00018350255,0.00048899447,0.0009370971,0.0033800537],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.99009305,0.004900205,0.00068696117,0.0007794467,0.0028299491,0.00071039214],"domain_scores_gemma":[0.932474,0.054014843,0.002343858,0.0051538055,0.005197607,0.0008157849],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013202412,0.00064298604,0.0007414013,0.00073616824,0.001235089,0.0053925617,0.0029348875,0.0016698899,0.004775442],"category_scores_gemma":[0.06894562,0.001193168,0.0009923276,0.0015227572,0.0021085397,0.0152776055,0.002715276,0.0036800087,0.0009406951],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003627791,0.0008381807,0.009471975,0.0018861312,0.0001519682,0.00038137563,0.0030706224,0.19176622,0.014007318,0.29328468,0.010829174,0.47394967],"study_design_scores_gemma":[0.00018550467,0.00040545277,0.0018383698,0.0004701527,0.00012195455,0.0003368953,0.004046931,0.6299748,0.03512015,0.27710524,0.050301507,0.00009298816],"about_ca_topic_score_codex":0.0059889844,"about_ca_topic_score_gemma":0.008384397,"teacher_disagreement_score":0.013202412,"about_ca_system_score_codex":0.0022308687,"about_ca_system_score_gemma":0.0056871986,"threshold_uncertainty_score":0.069821894},"labels":[],"label_agreement":null},{"id":"W4323851663","doi":"10.1002/smr.2548","title":"Combining object‐oriented metrics and centrality measures to predict faults in object‐oriented software: An empirical validation","year":2023,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada","keywords":"Centrality; Computer science; Software metric; Data mining; Object-oriented programming; Software; Software fault tolerance; Object (grammar); Fault (geology); Artificial intelligence; Software development; Machine learning; Software quality; Programming language; Mathematics","score_opus":0.033226737176487164,"score_gpt":0.3259724400237865,"score_spread":0.29274570284729934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323851663","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9925609,0.00013997978,0.006347441,0.00007266079,0.000015395373,0.00004829537,0.00036585185,0.00007054847,0.00037889887],"genre_scores_gemma":[0.99479973,0.00004965379,0.0042558596,0.00001896028,0.0000152462035,0.000059375743,0.0006342328,0.00002035051,0.00014641654],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99363756,0.003674457,0.00032973028,0.00087691366,0.0012367662,0.00024458038],"domain_scores_gemma":[0.7493601,0.21091343,0.010178173,0.0112413755,0.016520495,0.0017865513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016618649,0.0011250741,0.0006232799,0.0045993812,0.0005448883,0.0011807346,0.0015298737,0.0012822871,0.0012225023],"category_scores_gemma":[0.08240763,0.0004352794,0.0012019177,0.003134347,0.0010721356,0.0025031727,0.0013463182,0.0017523975,0.00057013525],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049390644,0.0015843165,0.9190753,0.00015872564,0.0006233841,0.00010323717,0.000548006,0.042080354,0.0009263286,0.00073121866,0.0011120968,0.032563154],"study_design_scores_gemma":[0.00012048493,0.0014859901,0.4216964,0.000106176594,0.00023615081,0.00012767209,0.00068390527,0.57011604,0.002403527,0.0019488316,0.0010171003,0.00005770046],"about_ca_topic_score_codex":0.0076579205,"about_ca_topic_score_gemma":0.0053787483,"teacher_disagreement_score":0.016618649,"about_ca_system_score_codex":0.00087773916,"about_ca_system_score_gemma":0.00078086293,"threshold_uncertainty_score":0.0878889},"labels":[],"label_agreement":null},{"id":"W4323925049","doi":"10.1007/s10664-022-10277-5","title":"Registered reports in software engineering","year":2023,"lang":"en","type":"editorial","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software engineering; Computer science; Engineering; Systems engineering","score_opus":0.027683895632145253,"score_gpt":0.2998760716857518,"score_spread":0.27219217605360657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323925049","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00004433598,0.0073862835,0.0011555975,0.0383524,0.93983257,0.000103404935,0.000117247124,0.0003423927,0.012665759],"genre_scores_gemma":[0.0014163454,0.018013798,0.0023014452,0.03740693,0.863208,0.00029015864,0.00035665085,0.0010912069,0.075915396],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.93733037,0.014283136,0.009447529,0.0035549945,0.033983506,0.0014005773],"domain_scores_gemma":[0.7015031,0.117403984,0.016086161,0.01748957,0.1319167,0.015600544],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.035975683,0.0033652834,0.0026170649,0.008670058,0.005011977,0.020694561,0.0053459504,0.017685637,0.044568412],"category_scores_gemma":[0.22360691,0.0014045018,0.0020250443,0.0066456976,0.0061459304,0.011708659,0.005214273,0.023398357,0.06401398],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011807922,0.000008886055,0.000017169563,0.00044542196,0.0000049728924,0.00006207055,0.00006234594,0.00002614291,0.000044624805,0.0030264524,0.9843912,0.0118990075],"study_design_scores_gemma":[0.0000052858395,0.0000060311972,0.00002291184,0.00035348273,0.0000038155017,0.0000627157,0.000030398362,0.000025060443,0.00003896135,0.0010956976,0.99834895,0.0000066554453],"about_ca_topic_score_codex":0.001186232,"about_ca_topic_score_gemma":0.0021090303,"teacher_disagreement_score":0.9640243,"about_ca_system_score_codex":0.0051859817,"about_ca_system_score_gemma":0.012480186,"threshold_uncertainty_score":0.19025987},"labels":[],"label_agreement":null},{"id":"W4324302602","doi":"10.48550/arxiv.2303.06233","title":"Model-Agnostic Syntactical Information for Pre-Trained Programming Language Models","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Automatic summarization; Abstract syntax tree; Source code; Language model; Security token; Transformer; Artificial intelligence; Scratch; Syntax; Programming language; Code (set theory); Natural language processing","score_opus":0.09852245320196315,"score_gpt":0.24330230736932135,"score_spread":0.1447798541673582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4324302602","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08585504,0.0014315506,0.8301272,0.0008987715,0.0003624698,0.00032693805,0.0030935358,0.07168156,0.0062229284],"genre_scores_gemma":[0.5409276,0.0007403505,0.4199628,0.0013845247,0.00014748298,0.0010077051,0.02046619,0.0042339168,0.011129404],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896836,0.0002510735,0.000064789085,0.00045861077,0.000151117,0.00010610196],"domain_scores_gemma":[0.9974366,0.0013332206,0.00012537635,0.00049408665,0.00052895275,0.0000818387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012462224,0.0025222471,0.0008479636,0.001208049,0.00042908208,0.0011549939,0.0027820214,0.0013346843,0.0046958956],"category_scores_gemma":[0.0060639246,0.0009336535,0.0018869146,0.0009010736,0.0005533071,0.0033844765,0.0014545068,0.0046015717,0.0040091584],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042742342,0.00035021166,0.0041485117,0.00060587394,0.00026894678,0.00024447186,0.0002291973,0.3783984,0.023603648,0.0039012157,0.025396168,0.5624259],"study_design_scores_gemma":[0.00002202371,0.000063009546,0.00037232618,0.00003059078,0.000049776154,0.000040162075,0.000029020692,0.9864125,0.007250076,0.0032091916,0.0025026028,0.000018731378],"about_ca_topic_score_codex":0.0067825625,"about_ca_topic_score_gemma":0.013920826,"teacher_disagreement_score":0.0067825625,"about_ca_system_score_codex":0.00145101,"about_ca_system_score_gemma":0.0019233779,"threshold_uncertainty_score":0.01570934},"labels":[],"label_agreement":null},{"id":"W4324374545","doi":"10.1007/s10664-022-10260-0","title":"Evaluating ensemble imputation in software effort estimation","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Imputation (statistics); Computer science; Missing data; Statistics; Data mining; Regression; Mathematics; Machine learning","score_opus":0.054878849205720864,"score_gpt":0.3653485525550296,"score_spread":0.3104697033493088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4324374545","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21975522,0.0016915279,0.77241814,0.0005957004,0.00032643866,0.00010366924,0.00081488857,0.0018583842,0.002436058],"genre_scores_gemma":[0.7629848,0.0003597117,0.23007214,0.00032172966,0.00032532163,0.00019310316,0.0034031237,0.00024276107,0.0020973256],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9901484,0.006945865,0.00042222277,0.0011519665,0.0009559924,0.00037550958],"domain_scores_gemma":[0.8882717,0.093457475,0.0022697921,0.010106132,0.0048724324,0.0010224406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022441618,0.0009664868,0.0028661923,0.0016183393,0.0007777405,0.0015657732,0.0027468514,0.002718533,0.00253115],"category_scores_gemma":[0.08030682,0.00080683484,0.0014426069,0.0023163997,0.0005754908,0.003027127,0.002319292,0.0027768116,0.0008424649],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013491473,0.00061931554,0.045051377,0.0001794758,0.0013074942,0.00015926201,0.00020130954,0.6885693,0.0006572745,0.0071832957,0.0067350725,0.24798761],"study_design_scores_gemma":[0.00003363999,0.00010511921,0.002210315,0.000024650653,0.00006987812,0.000028324792,0.000028373594,0.9897411,0.00038840825,0.006947999,0.00041135112,0.000010835768],"about_ca_topic_score_codex":0.003716601,"about_ca_topic_score_gemma":0.004213605,"teacher_disagreement_score":0.022441618,"about_ca_system_score_codex":0.00058110745,"about_ca_system_score_gemma":0.0012490953,"threshold_uncertainty_score":0.118683994},"labels":[],"label_agreement":null},{"id":"W4360764539","doi":"10.1109/icmla55696.2022.00173","title":"VDGraph2Vec: Vulnerability Detection in Assembly Code using Message Passing Neural Networks","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Canada Research Chairs","keywords":"Computer science; Vulnerability (computing); Malware; Deep learning; Reverse engineering; Software; Artificial intelligence; Benchmark (surveying); Task (project management); Machine learning; Process (computing); Vulnerability management; Artificial neural network; Software security assurance; Code (set theory); Software engineering; Static program analysis; Computer security; Vulnerability assessment; Software development; Information security; Programming language; Engineering","score_opus":0.030222659632276615,"score_gpt":0.29015073434065447,"score_spread":0.25992807470837787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4360764539","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3513584,0.0020013568,0.58607745,0.0015523914,0.0006684254,0.00027725683,0.008091427,0.043707494,0.006265766],"genre_scores_gemma":[0.7430179,0.0007881001,0.22213,0.0006001734,0.000122001235,0.0002859996,0.020778002,0.0009661697,0.011311769],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99978226,0.0000349697,0.000012028735,0.00007079658,0.000060093054,0.00003983569],"domain_scores_gemma":[0.999572,0.0001581725,0.000052584393,0.000063686726,0.00012997992,0.000023556768],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00030231642,0.00162575,0.00045819933,0.0012487016,0.00032287854,0.0005609179,0.0010950378,0.0008691947,0.0016535973],"category_scores_gemma":[0.0013930963,0.00044966012,0.0007339001,0.0008272774,0.0004028161,0.0011086382,0.00065371866,0.0011966726,0.000956258],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034002203,0.00030359568,0.010350401,0.0003286605,0.00025040397,0.0004854057,0.00019260723,0.411212,0.01732499,0.004921623,0.048583787,0.50570655],"study_design_scores_gemma":[0.000009424391,0.00004905509,0.00079208065,0.000010078105,0.000015077032,0.0000526794,0.000016548367,0.9905019,0.0048352797,0.0018481727,0.0018591817,0.00001049077],"about_ca_topic_score_codex":0.016569719,"about_ca_topic_score_gemma":0.027561447,"teacher_disagreement_score":0.016569719,"about_ca_system_score_codex":0.0008097987,"about_ca_system_score_gemma":0.00091395725,"threshold_uncertainty_score":0.032946587},"labels":[],"label_agreement":null},{"id":"W4360885093","doi":"10.5281/zenodo.7765615","title":"Predicting the change propagation of code clones at a pull request level","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Code (set theory); Computer science; Programming language","score_opus":0.1115727864960143,"score_gpt":0.2731884045473224,"score_spread":0.16161561805130809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4360885093","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8876876,0.0003369248,0.098420486,0.0005677584,0.00025005694,0.00017176954,0.0028612236,0.008109857,0.0015943007],"genre_scores_gemma":[0.9356666,0.00014675497,0.056924675,0.00011496871,0.00011376058,0.00008896115,0.0042908853,0.0004859917,0.0021673057],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99723816,0.00033718435,0.00014378219,0.0007406099,0.0013096032,0.00023073009],"domain_scores_gemma":[0.97660017,0.0118240025,0.0031725857,0.0031389748,0.004367759,0.0008965076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001768489,0.0010368954,0.0007862838,0.0022865988,0.0004850282,0.0010761316,0.0009598169,0.0016349939,0.0015732739],"category_scores_gemma":[0.019470321,0.00057354337,0.0010054561,0.0015860625,0.00046841492,0.0014369492,0.0008712199,0.001526223,0.001108831],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012535504,0.00077908393,0.40748933,0.0005218992,0.00033482234,0.00087747845,0.0004185591,0.26881614,0.09863539,0.001903278,0.008508287,0.21046236],"study_design_scores_gemma":[0.000024664743,0.00029194748,0.05003151,0.00001158881,0.00005166194,0.00020460246,0.000049179856,0.9203057,0.0269238,0.0011745646,0.0008936464,0.000037151945],"about_ca_topic_score_codex":0.0046121106,"about_ca_topic_score_gemma":0.0055258176,"teacher_disagreement_score":0.0046121106,"about_ca_system_score_codex":0.0006218813,"about_ca_system_score_gemma":0.0008643648,"threshold_uncertainty_score":0.009352744},"labels":[],"label_agreement":null},{"id":"W4361193297","doi":"10.48550/arxiv.2303.14542","title":"Combining Contexts from Multiple Sources for Documentation-Specific Code Example Generation","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates","keywords":"Documentation; Internal documentation; Computer science; Code (set theory); Programming language; Compiler; Source code; Software documentation; Unit testing; Software; Software development; Set (abstract data type); Software development process; Software construction","score_opus":0.19255112669707944,"score_gpt":0.2377315074623969,"score_spread":0.04518038076531747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4361193297","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36519238,0.0024246844,0.5834044,0.0014451608,0.00031678195,0.0010646044,0.002521201,0.03358493,0.010045898],"genre_scores_gemma":[0.5080509,0.00046036066,0.47921792,0.0004199652,0.00006730002,0.0006266254,0.0055916775,0.0021422328,0.0034230142],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979219,0.0010413831,0.00012395623,0.00045881377,0.00035279666,0.000101137426],"domain_scores_gemma":[0.98989797,0.006160955,0.000411847,0.0021380354,0.0011698665,0.00022131135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002382135,0.0012217241,0.00054886466,0.0016922768,0.00047602103,0.001081826,0.0013156203,0.0011686594,0.0024373643],"category_scores_gemma":[0.020587578,0.00056499505,0.00086935295,0.001064238,0.0005165846,0.0023232994,0.0023946774,0.0017396732,0.0016489641],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007809347,0.000731328,0.03756999,0.0018015108,0.00019150763,0.0014446211,0.0027855455,0.038691882,0.026380595,0.0061879056,0.020647164,0.862787],"study_design_scores_gemma":[0.00026081625,0.00071763754,0.012685388,0.00072389084,0.00027083795,0.0016544005,0.0011562161,0.83554834,0.05787905,0.026484745,0.062470384,0.00014821389],"about_ca_topic_score_codex":0.0013475533,"about_ca_topic_score_gemma":0.0045607896,"teacher_disagreement_score":0.0024373643,"about_ca_system_score_codex":0.00045102293,"about_ca_system_score_gemma":0.0011950248,"threshold_uncertainty_score":0.012598038},"labels":[],"label_agreement":null},{"id":"W4361275045","doi":"10.1007/s11219-023-09621-9","title":"Machine learning application development: practitioners’ insights","year":2023,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada","keywords":"Best practice; Leverage (statistics); Computer science; Engineering management; Process (computing); Software development process; Software development; Quality (philosophy); Software; Knowledge management; Software engineering; Data science; Engineering; Artificial intelligence","score_opus":0.0405552949918816,"score_gpt":0.3248114528857969,"score_spread":0.2842561578939153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4361275045","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1256034,0.03158413,0.16373943,0.5472318,0.0013643757,0.00009472644,0.00013683864,0.00034565103,0.12989968],"genre_scores_gemma":[0.9048167,0.019053675,0.044007555,0.019202244,0.0012849413,0.00008801453,0.000087576635,0.00012416749,0.011335108],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9937441,0.0030052215,0.0002561725,0.0004860827,0.0021923606,0.0003160739],"domain_scores_gemma":[0.93944794,0.04315346,0.0009997446,0.0020713096,0.012425234,0.0019022131],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01790935,0.00038630015,0.00029625694,0.0025672794,0.0011554825,0.004738406,0.0011819782,0.0027773925,0.0044612503],"category_scores_gemma":[0.039952822,0.00031023848,0.00021273071,0.0018198427,0.0029635644,0.008907821,0.002518629,0.0049652806,0.00088392093],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012201308,0.0006159809,0.01568779,0.0010952855,0.000044302327,0.00071769155,0.030418277,0.0036047988,0.0033776027,0.4408398,0.07277441,0.43070215],"study_design_scores_gemma":[0.00009772564,0.00034098106,0.010643829,0.0024401096,0.00006204987,0.0016429134,0.051803984,0.030603547,0.0058462494,0.49496153,0.40147102,0.0000860203],"about_ca_topic_score_codex":0.0016018257,"about_ca_topic_score_gemma":0.0031430929,"teacher_disagreement_score":0.01790935,"about_ca_system_score_codex":0.0018176247,"about_ca_system_score_gemma":0.003185698,"threshold_uncertainty_score":0.09471482},"labels":[],"label_agreement":null},{"id":"W4362505496","doi":"10.1145/3578245.3584695","title":"Software Mining -- Investigating Correlation between Source Code Features and Michrobenchmark's Steady State","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"","keywords":"Source code; Computer science; Steady state (chemistry); Software; Code (set theory); Open source; Open source software; Programming language; Chemistry","score_opus":0.02684432476691485,"score_gpt":0.27211056965926644,"score_spread":0.2452662448923516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362505496","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9862499,0.00022324301,0.010461964,0.00007683491,0.00000871469,0.000047892598,0.0013083747,0.00055322936,0.0010698745],"genre_scores_gemma":[0.9909559,0.000079524594,0.006665467,0.0000140165375,0.000007109939,0.00003612567,0.0019569923,0.000039486873,0.0002453556],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982541,0.00022624487,0.00028859798,0.0004996489,0.00059998885,0.00013135973],"domain_scores_gemma":[0.9683427,0.017042048,0.007780673,0.0020203965,0.004253552,0.0005605587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018343654,0.00039468528,0.0004168983,0.0039064987,0.00036735274,0.0008596649,0.0005358107,0.00030731913,0.0005216385],"category_scores_gemma":[0.017953392,0.0001975405,0.00046192852,0.0035635468,0.00033055618,0.0010111324,0.0004937916,0.000503072,0.00032010808],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015944039,0.00019146026,0.9093622,0.0001592398,0.00012386432,0.00018096933,0.00026506695,0.003794395,0.009237979,0.00035908146,0.00064141187,0.07552488],"study_design_scores_gemma":[0.000008840741,0.00037387386,0.9199088,0.000036399175,0.000077005774,0.0007872242,0.00046102123,0.058349915,0.01725058,0.0011236067,0.0015878817,0.00003484587],"about_ca_topic_score_codex":0.0015298758,"about_ca_topic_score_gemma":0.0022008282,"teacher_disagreement_score":0.0039064987,"about_ca_system_score_codex":0.00034831176,"about_ca_system_score_gemma":0.00054317265,"threshold_uncertainty_score":0.009701133},"labels":[],"label_agreement":null},{"id":"W4362513832","doi":"10.1007/978-3-031-29786-1_6","title":"Scope Determined (D) and Scope Determining (G) Requirements: A New Categorization of Functional Requirements","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Scope (computer science); Computer science; Requirements management; Requirements analysis; Requirements engineering; Context (archaeology); Functional requirement; Non-functional requirement; Software engineering; System requirements specification; Statement of work; Software requirements specification; User requirements document; Requirement; Systems engineering; Agile software development; Risk analysis (engineering); Software; Software development; Engineering; Software design; Programming language; Business","score_opus":0.0553492677093169,"score_gpt":0.29020898054831173,"score_spread":0.23485971283899484,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362513832","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02607723,0.0036524932,0.88565147,0.0018384604,0.0002440574,0.00064527604,0.00048741448,0.0011111051,0.0802925],"genre_scores_gemma":[0.2492833,0.002439506,0.7294468,0.00074643997,0.00022123361,0.00082970964,0.0011787373,0.00061754644,0.015236671],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9932106,0.0017901938,0.00084095827,0.0006512102,0.0031344777,0.00037250883],"domain_scores_gemma":[0.99007726,0.0051185326,0.0009910879,0.0014468009,0.0019007984,0.00046544723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040506986,0.0010337756,0.000705613,0.0055579655,0.0010030527,0.0053326506,0.002000124,0.0022513813,0.004112612],"category_scores_gemma":[0.011944778,0.0005659579,0.0016451918,0.004052342,0.004765164,0.007384309,0.0021670938,0.0025616481,0.0014987084],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000098797165,0.00013046575,0.0028132845,0.00046403354,0.000036097837,0.00047217667,0.005262947,0.0027236557,0.01016203,0.73372936,0.0066132275,0.23749392],"study_design_scores_gemma":[0.0001027221,0.0005686576,0.008613703,0.001108431,0.00017354013,0.0037460958,0.0047019473,0.049409296,0.007192854,0.74656636,0.17758286,0.00023340843],"about_ca_topic_score_codex":0.0035889477,"about_ca_topic_score_gemma":0.0029304551,"teacher_disagreement_score":0.0055579655,"about_ca_system_score_codex":0.0019821443,"about_ca_system_score_gemma":0.0031424,"threshold_uncertainty_score":0.021422386},"labels":[],"label_agreement":null},{"id":"W4362514913","doi":"10.1007/978-3-031-29786-1_7","title":"Using Language Models for Enhancing the Completeness of Natural-Language Requirements","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Terminology; Completeness (order theory); Natural language; Context (archaeology); Principal (computer security); Measure (data warehouse); Artificial intelligence; Noise (video); Filter (signal processing); Language model; Natural language processing; Data mining; Linguistics; Computer security","score_opus":0.06258945467229786,"score_gpt":0.32607047861426136,"score_spread":0.2634810239419635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362514913","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01222097,0.000099376295,0.98245054,0.00040383544,0.000035004097,0.00015876061,0.00014160325,0.0014746793,0.0030152032],"genre_scores_gemma":[0.22123197,0.00026965444,0.7730824,0.00026960173,0.00006890587,0.00034003262,0.0010057975,0.0010406204,0.002691007],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9881837,0.005663488,0.0007562116,0.0009430028,0.0039658537,0.00048777924],"domain_scores_gemma":[0.95708305,0.02885012,0.0020506997,0.0077527715,0.0039283605,0.0003350572],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008380263,0.0011990629,0.0009469431,0.0016104969,0.0007461868,0.002864759,0.0021052235,0.0012873355,0.0038408726],"category_scores_gemma":[0.043977857,0.0014437374,0.0031663477,0.0011014784,0.0016546722,0.008491837,0.0039982456,0.0034472651,0.0010504278],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005351441,0.00058555615,0.0021848178,0.0011788227,0.00025495663,0.00052869064,0.0018430798,0.30131754,0.029028378,0.38611677,0.00815222,0.26827398],"study_design_scores_gemma":[0.0001113426,0.00014629778,0.00020073613,0.0001544383,0.00012323221,0.00018420319,0.0002038088,0.75154334,0.021580309,0.21300699,0.012688886,0.00005648146],"about_ca_topic_score_codex":0.0026223874,"about_ca_topic_score_gemma":0.00568363,"teacher_disagreement_score":0.008380263,"about_ca_system_score_codex":0.001396157,"about_ca_system_score_gemma":0.0032329164,"threshold_uncertainty_score":0.04431963},"labels":[],"label_agreement":null},{"id":"W4362584426","doi":"10.1007/s10664-022-10282-8","title":"Towards a change taxonomy for machine learning pipelines","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; York University; Ansys (Canada); Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Metadata; Replicate; Pipeline (software); Taxonomy (biology); Implementation; Data science; Fork (system call); Source code; Information retrieval; World Wide Web; Software engineering; Programming language; Ecology","score_opus":0.10899159167165356,"score_gpt":0.316721378863043,"score_spread":0.20772978719138946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362584426","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008185324,0.0003548082,0.97752225,0.0018878196,0.00011112942,0.00058459,0.0006566226,0.0028856823,0.00781175],"genre_scores_gemma":[0.10375084,0.00045673255,0.8860087,0.0004560165,0.000096777956,0.0005365134,0.0021338311,0.0006065238,0.005954081],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9932827,0.0014674774,0.0010464598,0.0015697334,0.0020784335,0.0005551683],"domain_scores_gemma":[0.9727721,0.007659595,0.0016879733,0.0071619023,0.009246711,0.0014717394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062137702,0.0012458237,0.0011757078,0.008686212,0.0035130822,0.00811958,0.003956741,0.004507366,0.007276356],"category_scores_gemma":[0.025822083,0.0015129816,0.003594839,0.004998873,0.0050920458,0.021204907,0.005503446,0.0073338407,0.004135376],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001382553,0.00028867414,0.012286085,0.00054337154,0.00007400111,0.00047947798,0.0029746436,0.016683828,0.0027977936,0.73124546,0.013149568,0.21933885],"study_design_scores_gemma":[0.00004303426,0.0001646883,0.0025945865,0.000341714,0.00009271681,0.00051873765,0.0011612199,0.23625317,0.0030503205,0.6670912,0.08860309,0.000085554464],"about_ca_topic_score_codex":0.011677766,"about_ca_topic_score_gemma":0.010235367,"teacher_disagreement_score":0.011677766,"about_ca_system_score_codex":0.0031418141,"about_ca_system_score_gemma":0.0054108123,"threshold_uncertainty_score":0.032861948},"labels":[],"label_agreement":null},{"id":"W4362597934","doi":"10.48550/arxiv.2304.00078","title":"A Meta-Summary of Challenges in Building Products with ML Components -- Collecting Experiences from 4758+ Practitioners","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Carnegie Mellon University; Natural Sciences and Engineering Research Council of Canada; U.S. Department of Defense; National Science Foundation","keywords":"Systematic review; Resource (disambiguation); Field (mathematics); Knowledge management; Data science; Computer science; Medical education; Engineering ethics; Engineering; Political science; MEDLINE; Medicine","score_opus":0.30875622643774925,"score_gpt":0.24728463018895747,"score_spread":0.06147159624879178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362597934","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0721114,0.8423293,0.02221429,0.017362978,0.0019741459,0.002228001,0.02342915,0.0007159456,0.017634783],"genre_scores_gemma":[0.32105643,0.5964469,0.04084673,0.010399162,0.00081136986,0.006327437,0.019143227,0.0008246593,0.004144119],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9606049,0.01597468,0.013681497,0.0022252523,0.0067682704,0.00074545504],"domain_scores_gemma":[0.70911944,0.2329612,0.018962536,0.010198355,0.026876004,0.0018824337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.041612267,0.0014481181,0.002072829,0.029016705,0.0018113677,0.0066131246,0.0017154391,0.0021591051,0.006413188],"category_scores_gemma":[0.20379765,0.0013965779,0.005855612,0.025145326,0.0012434728,0.009201489,0.0046071988,0.0022074003,0.0015028327],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082235265,0.00011721547,0.022824086,0.41432774,0.009913732,0.0008883844,0.036264617,0.0009820735,0.0041066646,0.005051024,0.048448894,0.4562533],"study_design_scores_gemma":[0.00017862426,0.00083872024,0.028006857,0.56657135,0.027336258,0.0017664545,0.021087868,0.00043616723,0.0036705555,0.00549113,0.34438556,0.0002303713],"about_ca_topic_score_codex":0.0042192847,"about_ca_topic_score_gemma":0.010124661,"teacher_disagreement_score":0.041612267,"about_ca_system_score_codex":0.004548724,"about_ca_system_score_gemma":0.010210446,"threshold_uncertainty_score":0.22006935},"labels":[],"label_agreement":null},{"id":"W4362603997","doi":"10.1016/j.infsof.2023.107218","title":"Dev2vec: Representing domain expertise of developers in an embedding space","year":2023,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Embedding; Computer science; Domain (mathematical analysis); Space (punctuation); Software engineering; Data science; Software; World Wide Web; Information retrieval; Knowledge management; Artificial intelligence; Programming language; Mathematics","score_opus":0.014334294860793103,"score_gpt":0.2854876674304545,"score_spread":0.2711533725696614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362603997","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34474847,0.00347838,0.5350392,0.002120835,0.0014032564,0.0003472124,0.06629265,0.028834354,0.017735593],"genre_scores_gemma":[0.73162496,0.001059688,0.16922571,0.00040362857,0.00037774502,0.00040773491,0.08439659,0.0010997641,0.011404079],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991346,0.00018654615,0.00006015501,0.00028621245,0.00021739514,0.0001151667],"domain_scores_gemma":[0.9985025,0.00054658524,0.000110898734,0.00024465047,0.0004912446,0.0001041096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006434666,0.001461121,0.0004920095,0.0032444734,0.00036885694,0.0010272538,0.0005734811,0.0009524236,0.0042733992],"category_scores_gemma":[0.0037568873,0.00027279626,0.0006059397,0.0023204111,0.00034674729,0.0018722009,0.0012246496,0.0010779045,0.0025204602],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085879426,0.00066049374,0.03256924,0.0008810565,0.00030383203,0.00038040118,0.00076925056,0.029789958,0.020372378,0.008071357,0.20517194,0.7001713],"study_design_scores_gemma":[0.00012846268,0.0005163457,0.024799768,0.00026101526,0.00022890534,0.0007657869,0.00088454824,0.8328425,0.026833624,0.025458481,0.08717135,0.000109276465],"about_ca_topic_score_codex":0.005562585,"about_ca_topic_score_gemma":0.01115869,"teacher_disagreement_score":0.005562585,"about_ca_system_score_codex":0.00044539903,"about_ca_system_score_gemma":0.0008301522,"threshold_uncertainty_score":0.014295876},"labels":[],"label_agreement":null},{"id":"W4362657725","doi":"10.1371/journal.pone.0283838","title":"Prioritizing tasks in software development: A systematic literature review","year":2023,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Software Engineering Research","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Huawei Technologies","keywords":"Computer science; Prioritization; Systematic review; Task (project management); Requirement prioritization; Ranking (information retrieval); Software; Domain (mathematical analysis); Software development; Data science; Field (mathematics); Risk analysis (engineering); Software engineering; Process management; Software construction; Artificial intelligence; Systems engineering; Engineering; Medicine; MEDLINE","score_opus":0.0481854037844086,"score_gpt":0.2629992655593275,"score_spread":0.2148138617749189,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362657725","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019715724,0.99265516,0.0017431027,0.0008935186,0.00031361845,0.0015466698,0.00046126192,0.000029113588,0.0003860326],"genre_scores_gemma":[0.026791403,0.9610528,0.0067590773,0.0011252491,0.00022011451,0.0033505808,0.0005249732,0.000023902581,0.00015195571],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.94714856,0.02126689,0.018018745,0.0029982263,0.009678669,0.00088881917],"domain_scores_gemma":[0.82016677,0.13183579,0.022395208,0.003663146,0.020328922,0.0016101488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04347318,0.0025934093,0.00836751,0.03641022,0.001831817,0.004481221,0.002666838,0.0029051157,0.002659659],"category_scores_gemma":[0.17708509,0.0020164992,0.009959069,0.024691677,0.0018008858,0.007686356,0.004029556,0.0023898226,0.00044277275],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018956434,0.00004783325,0.0016012563,0.8885678,0.005869927,0.00019550289,0.0007413712,0.00028459067,0.00026632115,0.00058676285,0.0023422677,0.09930664],"study_design_scores_gemma":[0.00016936185,0.00023880591,0.002997732,0.9536251,0.024380863,0.0004673259,0.0007923173,0.00021379921,0.00027614768,0.0009077983,0.01586332,0.00006740451],"about_ca_topic_score_codex":0.008265697,"about_ca_topic_score_gemma":0.024419665,"teacher_disagreement_score":0.04347318,"about_ca_system_score_codex":0.007936184,"about_ca_system_score_gemma":0.04002806,"threshold_uncertainty_score":0.22991091},"labels":[],"label_agreement":null},{"id":"W4364321651","doi":"10.1109/tse.2023.3265855","title":"Identifying Concepts in Software Projects","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Documentation; Software engineering; Domain (mathematical analysis); Software project management; Software development; Software; Domain analysis; Data science; Software construction; Programming language","score_opus":0.03530785764332117,"score_gpt":0.2955445288280574,"score_spread":0.2602366711847362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4364321651","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38433805,0.0027643899,0.5825015,0.0011791515,0.00018807875,0.00081111636,0.0067994236,0.0025884062,0.01882985],"genre_scores_gemma":[0.3974297,0.0011610824,0.58830225,0.00017157746,0.00004668442,0.0007902655,0.0087039005,0.00039060012,0.0030038897],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.995284,0.0012171986,0.00058513955,0.0011087258,0.0016118664,0.00019311512],"domain_scores_gemma":[0.9835096,0.008421813,0.0023099163,0.0014914752,0.003724908,0.0005422031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032699876,0.0007465834,0.00040846676,0.0142196175,0.0015535436,0.0029866546,0.0010002928,0.0012665759,0.0022482718],"category_scores_gemma":[0.03015948,0.0004807557,0.0007231347,0.011668959,0.0010598444,0.008359853,0.004711513,0.001297858,0.0008022797],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023845003,0.0002323333,0.09286644,0.0014793681,0.000109467386,0.0010492381,0.01567258,0.0062195812,0.010596889,0.08616657,0.01172877,0.7736403],"study_design_scores_gemma":[0.00012900801,0.00039715122,0.10532776,0.002017922,0.00030640032,0.004391159,0.024560861,0.14678283,0.029036734,0.32280064,0.3639205,0.00032906773],"about_ca_topic_score_codex":0.00509015,"about_ca_topic_score_gemma":0.004294871,"teacher_disagreement_score":0.0142196175,"about_ca_system_score_codex":0.001341047,"about_ca_system_score_gemma":0.002744616,"threshold_uncertainty_score":0.017293513},"labels":[],"label_agreement":null},{"id":"W4364355514","doi":"10.1016/j.infsof.2023.107229","title":"To group or not to group? Group sizes for requirements elicitation","year":2023,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Requirements elicitation; Viewpoints; Requirements engineering; Requirements management; Expert elicitation; Requirements analysis; Computer science; Software; Software engineering; Knowledge management; Engineering; Mathematics","score_opus":0.030302194558049594,"score_gpt":0.31377786682737885,"score_spread":0.28347567226932924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4364355514","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3523704,0.0018896717,0.5600195,0.009250319,0.0019414257,0.019334733,0.0020590886,0.003371005,0.0497639],"genre_scores_gemma":[0.6966865,0.00019119825,0.28523663,0.00077356986,0.00015875247,0.013751935,0.00042025922,0.00027074214,0.002510378],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9250286,0.057913553,0.003584897,0.0034554354,0.009169795,0.0008477062],"domain_scores_gemma":[0.5509807,0.41081485,0.006517271,0.019505594,0.009616532,0.0025651057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07128897,0.0006877715,0.0014617731,0.0017973996,0.0015795908,0.0017984256,0.0025846027,0.0016149329,0.01687128],"category_scores_gemma":[0.3057353,0.0006143171,0.00083515036,0.0011029642,0.0017894402,0.005155875,0.002717777,0.0018196995,0.0022779298],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.01449976,0.0018526958,0.008871221,0.0018544652,0.00048026888,0.00015278022,0.008562908,0.007043003,0.015919631,0.047336116,0.024194816,0.8692323],"study_design_scores_gemma":[0.025759296,0.02674361,0.09942034,0.0036069097,0.0022122904,0.0011362175,0.025952568,0.22040051,0.055063315,0.43158668,0.10705825,0.0010600236],"about_ca_topic_score_codex":0.00089065183,"about_ca_topic_score_gemma":0.0019737226,"teacher_disagreement_score":0.07128897,"about_ca_system_score_codex":0.0016708337,"about_ca_system_score_gemma":0.0020579689,"threshold_uncertainty_score":0.3770166},"labels":[],"label_agreement":null},{"id":"W4365456905","doi":"10.1016/j.infsof.2023.107222","title":"A mapping study of language features improving object-oriented design patterns","year":2023,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Université de Montréal; Concordia University","funders":"Fonds de recherche du Québec – Nature et technologies","keywords":"Software design pattern; Computer science; Design pattern; Implementation; Object-oriented design; Software design; Categorization; Software engineering; Structural pattern; Visitor pattern; Software; Object-oriented programming; Programming language; Software development; Artificial intelligence","score_opus":0.011371571548867558,"score_gpt":0.25075650824291007,"score_spread":0.2393849366940425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4365456905","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75271326,0.00018699888,0.23488052,0.00028921993,0.000033115226,0.0002502701,0.00018538606,0.00093769666,0.010523465],"genre_scores_gemma":[0.8542611,0.00012844882,0.14203405,0.00004651779,0.000006803585,0.000083812796,0.00019383819,0.00019376722,0.0030516398],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.998904,0.0004023364,0.00007831308,0.00016831451,0.0003742542,0.00007270243],"domain_scores_gemma":[0.99112797,0.005165158,0.00073905307,0.0012216492,0.0016124249,0.00013378778],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0013224506,0.0003307494,0.00025934115,0.0011934699,0.0003970766,0.0012090632,0.0006150327,0.00040921604,0.00359109],"category_scores_gemma":[0.011261778,0.00026793458,0.00047515632,0.0014572089,0.00043997815,0.0026294882,0.0005147745,0.0004924273,0.00041814047],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008123863,0.0017120863,0.05767097,0.0008178456,0.00010619909,0.000535202,0.0068227486,0.018573632,0.08077253,0.05834487,0.0021038605,0.7717277],"study_design_scores_gemma":[0.00047606384,0.0043216315,0.09893289,0.00032569768,0.00094849034,0.0021056177,0.011701028,0.579372,0.17646219,0.08802489,0.037169836,0.00015969899],"about_ca_topic_score_codex":0.001581951,"about_ca_topic_score_gemma":0.001668396,"teacher_disagreement_score":0.99867755,"about_ca_system_score_codex":0.00043650722,"about_ca_system_score_gemma":0.0008720949,"threshold_uncertainty_score":0.012013376},"labels":[],"label_agreement":null},{"id":"W4366149280","doi":"10.1016/j.procs.2023.03.077","title":"Maintenance Cost of Software Ecosystem Updates","year":2023,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Software; Software maintenance; Software development; Ecosystem; Software engineering; Software analytics; Software construction; Operating system; Ecology","score_opus":0.016773421547049524,"score_gpt":0.2577916063472722,"score_spread":0.2410181848002227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366149280","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9264377,0.0015162546,0.059237655,0.0007997404,0.0001165719,0.00013101583,0.0013427553,0.0009022784,0.009515896],"genre_scores_gemma":[0.98937464,0.00032326102,0.008137965,0.000025753414,0.000013469877,0.000031425032,0.0005496926,0.000085394146,0.0014584071],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9976361,0.000448653,0.00018847818,0.00034528112,0.0010974377,0.0002839914],"domain_scores_gemma":[0.9880372,0.0071523357,0.0014373381,0.0011942144,0.0018155071,0.00036344162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024934225,0.00054894073,0.0005045333,0.0023950306,0.0005691856,0.0018362699,0.0014609569,0.0008985949,0.0023741426],"category_scores_gemma":[0.022113103,0.00035965673,0.0006096584,0.0017139185,0.0005976571,0.002832546,0.00090765517,0.000722814,0.00028203998],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051342184,0.00021359583,0.05106789,0.00034610956,0.00015279498,0.000850568,0.00038975204,0.78688365,0.007581998,0.022137497,0.0044949283,0.12536784],"study_design_scores_gemma":[0.000018829085,0.00020963908,0.015684092,0.000042354543,0.0001265991,0.00056975434,0.00025338284,0.9715897,0.00292496,0.0064573265,0.00208729,0.000036223475],"about_ca_topic_score_codex":0.0120795015,"about_ca_topic_score_gemma":0.0088559305,"teacher_disagreement_score":0.0120795015,"about_ca_system_score_codex":0.0025793407,"about_ca_system_score_gemma":0.0011156311,"threshold_uncertainty_score":0.024018407},"labels":[],"label_agreement":null},{"id":"W4366721896","doi":"10.1002/9781119880929.ch1","title":"Software Fault Localization: an Overview of Research, Techniques, and Tools","year":2023,"lang":"en","type":"other","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of Standards and Technology; Technische Universiteit Delft; McMaster University; Purdue University; TU Graz, Internationale Beziehungen und Mobilitätsprogramme; Brown University; Carnegie Mellon University; National University of Singapore; University of Cambridge; Yale University","keywords":"Profiling (computer programming); Computer science; Program slicing; Software; Debugging; Data mining; Slicing; Data science; Artificial intelligence; Machine learning; Programming language; World Wide Web","score_opus":0.14135512187788293,"score_gpt":0.40266536064629527,"score_spread":0.2613102387684123,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366721896","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041391635,0.7525803,0.15320897,0.005382794,0.0014817452,0.00020772112,0.000854418,0.004081874,0.07806303],"genre_scores_gemma":[0.023422979,0.7857004,0.15130036,0.0016135526,0.001575775,0.00026343588,0.0016824665,0.0009917667,0.033449207],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99804807,0.00026671373,0.00019919776,0.00029169224,0.0010862163,0.00010805579],"domain_scores_gemma":[0.9976071,0.0012525349,0.00018359478,0.00018672798,0.0006713142,0.00009864096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013880573,0.001553821,0.0009850633,0.008867134,0.00079452177,0.004878116,0.0013023013,0.0019466034,0.008158432],"category_scores_gemma":[0.0036719108,0.0013521815,0.00086437684,0.011027567,0.0013789934,0.00782307,0.0014653724,0.0025984775,0.0071049985],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004163558,0.000112660025,0.00082965114,0.004138857,0.00003383267,0.00012370758,0.0006139952,0.0030309719,0.0026794036,0.04850976,0.05202297,0.88786256],"study_design_scores_gemma":[0.000011659841,0.0001135606,0.0017778527,0.0039319443,0.00005728239,0.0013185459,0.0006675234,0.00582552,0.0033080387,0.050502688,0.9324082,0.000077104385],"about_ca_topic_score_codex":0.003141732,"about_ca_topic_score_gemma":0.0031943696,"teacher_disagreement_score":0.008867134,"about_ca_system_score_codex":0.0018162298,"about_ca_system_score_gemma":0.0019748735,"threshold_uncertainty_score":0.027292728},"labels":[],"label_agreement":null},{"id":"W4366825753","doi":"10.1145/3593802","title":"Predicting the Change Impact of Resolving Defects by Leveraging the Topics of Issue Reports in Open Source Software Systems","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Leverage (statistics); Data science; Open source; Metric (unit); Source code; Software bug; Change impact analysis; Eclipse; Data mining; Software; Information retrieval; Machine learning","score_opus":0.11775826133334773,"score_gpt":0.35457027254143647,"score_spread":0.23681201120808876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366825753","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9853727,0.0007073839,0.011665532,0.0002131856,0.000033874476,0.000055687036,0.0009966599,0.00047518694,0.0004798122],"genre_scores_gemma":[0.98601264,0.00030693188,0.009385046,0.00003800974,0.00006871526,0.00004728774,0.0037349402,0.000049107784,0.0003571521],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9984371,0.00038175102,0.00014833004,0.00039237822,0.0004743649,0.00016609501],"domain_scores_gemma":[0.9770641,0.015569097,0.004146455,0.0009568412,0.0017374066,0.000526124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036787344,0.0008970299,0.0005759919,0.0048333327,0.00038040458,0.0011808269,0.0007274609,0.0010347093,0.00037723835],"category_scores_gemma":[0.019941809,0.0003271177,0.0008907805,0.0031826862,0.00034351952,0.002265292,0.00084811467,0.0011193948,0.00035386696],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005603969,0.00064600335,0.744312,0.0004925133,0.00019798917,0.0004477146,0.00159155,0.074272655,0.007487223,0.0006131772,0.0044561857,0.16492274],"study_design_scores_gemma":[0.000027603128,0.0003207715,0.39321887,0.000058426576,0.00016250063,0.00030768022,0.0008108505,0.5966714,0.0044548623,0.0012037451,0.002720088,0.00004320094],"about_ca_topic_score_codex":0.0060377303,"about_ca_topic_score_gemma":0.005577978,"teacher_disagreement_score":0.0060377303,"about_ca_system_score_codex":0.0005857543,"about_ca_system_score_gemma":0.00048553868,"threshold_uncertainty_score":0.019455254},"labels":[],"label_agreement":null},{"id":"W4366990909","doi":"10.1007/s10664-023-10292-0","title":"Ranking code clones to support maintenance activities","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"clone (Java method); Code refactoring; Software maintenance; Commit; Computer science; Java; Code (set theory); Software evolution; Code reuse; Programming language; Software quality; Software system; Software development; Software; Database; Biology; Software construction; Genetics; Gene","score_opus":0.036875034545710464,"score_gpt":0.3080833253675743,"score_spread":0.2712082908218639,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366990909","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97322243,0.00075643236,0.020531673,0.00022252437,0.00004641489,0.00010598002,0.0007702247,0.0022452248,0.0020990598],"genre_scores_gemma":[0.9677697,0.00014689381,0.028804518,0.00004194994,0.000024094052,0.000036482557,0.0018235731,0.0001670525,0.0011857331],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99714965,0.0005845953,0.00019691678,0.00039122422,0.0014340053,0.00024355525],"domain_scores_gemma":[0.96125376,0.019230835,0.0051490488,0.0033992068,0.009621818,0.0013454024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017119055,0.0007126986,0.00050046586,0.0058109774,0.0005086603,0.0013753604,0.0008666951,0.001088642,0.0016891026],"category_scores_gemma":[0.032727152,0.00023923366,0.0005256946,0.0024538508,0.00031791348,0.0018504066,0.0005179546,0.00053505687,0.0006609417],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001470075,0.0007238979,0.42865133,0.00048137424,0.00022574578,0.0002882001,0.00044821863,0.015972309,0.0391219,0.0019387732,0.0066592377,0.5040189],"study_design_scores_gemma":[0.00043548076,0.003972293,0.39552104,0.00023063662,0.00089737476,0.0016328163,0.0014049537,0.4984501,0.07879375,0.008656189,0.009847863,0.00015741137],"about_ca_topic_score_codex":0.004368438,"about_ca_topic_score_gemma":0.010970183,"teacher_disagreement_score":0.0058109774,"about_ca_system_score_codex":0.0007971598,"about_ca_system_score_gemma":0.001615535,"threshold_uncertainty_score":0.009053528},"labels":[],"label_agreement":null},{"id":"W4367674071","doi":"10.1016/j.infsof.2023.107244","title":"A mixed method study of DevOps challenges","year":2023,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"DevOps; Computer science; Software engineering; Cloud computing; Orchestration; Empirical research; Information technology operations; Quality (philosophy); Data science; Engineering management; Engineering; Information technology; Software deployment","score_opus":0.0369279377940198,"score_gpt":0.2972357860131443,"score_spread":0.2603078482191245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4367674071","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9856938,0.00017022279,0.00685919,0.00018443953,0.000043099728,0.001523235,0.0002078253,0.000018991399,0.005299138],"genre_scores_gemma":[0.96505487,0.0003143074,0.016859418,0.0004796566,0.000055192057,0.0083711855,0.00036671918,0.000040121227,0.008458473],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9877851,0.008752934,0.00065712334,0.00090007356,0.0013383268,0.0005663228],"domain_scores_gemma":[0.9034808,0.07502502,0.0054879333,0.0050732708,0.00883215,0.0021008349],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014596172,0.00070685474,0.0006924588,0.0022844155,0.0035324893,0.0030733033,0.0017765387,0.0013591328,0.006405244],"category_scores_gemma":[0.048629723,0.0008348234,0.00044933,0.0019984161,0.0013956366,0.0031942697,0.002850422,0.0019030437,0.0011253142],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00574235,0.046690747,0.20655395,0.0035308902,0.00046863497,0.0017839726,0.39314649,0.0035696027,0.009945966,0.030409958,0.006663454,0.29149407],"study_design_scores_gemma":[0.003490845,0.040897753,0.21710761,0.0023667647,0.0005549708,0.0012873049,0.619248,0.017064735,0.0112904655,0.023311745,0.06289032,0.00048953894],"about_ca_topic_score_codex":0.0048599676,"about_ca_topic_score_gemma":0.009562524,"teacher_disagreement_score":0.014596172,"about_ca_system_score_codex":0.0030928012,"about_ca_system_score_gemma":0.004351126,"threshold_uncertainty_score":0.0771929},"labels":[],"label_agreement":null},{"id":"W4368408206","doi":"10.1145/3576841.3589626","title":"Automated Features and Requirements Identification for Improving CPS Software Reuse using Topic Modeling","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Software engineering; Documentation; Reuse; Software development; Software construction; Software system; Domain (mathematical analysis); Identification (biology); Software; Domain analysis; Programming language; Engineering","score_opus":0.06455301274145477,"score_gpt":0.3374883921143187,"score_spread":0.27293537937286394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4368408206","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06976782,0.00097413297,0.90732753,0.0006894408,0.000103711485,0.0007688265,0.0020377394,0.014477694,0.0038529742],"genre_scores_gemma":[0.33711964,0.00064060546,0.6479866,0.00024040489,0.000109053464,0.0008826835,0.009545499,0.0011070571,0.0023684385],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913891,0.0024753213,0.00083001325,0.002168326,0.002720598,0.00041673266],"domain_scores_gemma":[0.9777548,0.01250061,0.0025363334,0.0023477022,0.004403346,0.00045729842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048847036,0.0018815841,0.0013666827,0.011566695,0.0012334561,0.0027218226,0.0020941494,0.001978818,0.0016097333],"category_scores_gemma":[0.025272781,0.00079231977,0.0029750213,0.005145429,0.0005601357,0.0040355814,0.0023956217,0.0018615216,0.0021808676],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035028695,0.0008597436,0.033344265,0.0014668732,0.00026343967,0.0009776058,0.0035580741,0.024304146,0.044585362,0.0075942217,0.021880351,0.86081564],"study_design_scores_gemma":[0.00009083223,0.0002519475,0.017908191,0.00020126582,0.00040255795,0.0014329827,0.0014772429,0.9031218,0.033431284,0.013516016,0.02802538,0.00014053045],"about_ca_topic_score_codex":0.0096146455,"about_ca_topic_score_gemma":0.011883911,"teacher_disagreement_score":0.011566695,"about_ca_system_score_codex":0.0014081113,"about_ca_system_score_gemma":0.0030414115,"threshold_uncertainty_score":0.02583307},"labels":[],"label_agreement":null},{"id":"W4375824997","doi":"10.1007/978-3-031-28332-1_5","title":"A Two-Step Approach to Boost Neural Network Generalizability in Predicting Defective Software","year":2023,"lang":"en","type":"book-chapter","venue":"Advances in intelligent systems and computing","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Verafin (Canada)","funders":"","keywords":"Generalizability theory; Computer science; Artificial neural network; Machine learning; Software; Artificial intelligence; Software metric; Java; Quality (philosophy); Software quality; Data mining; Software engineering; Software development; Programming language; Statistics","score_opus":0.027673676825214814,"score_gpt":0.2816586271396522,"score_spread":0.25398495031443735,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4375824997","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06338208,0.0014575215,0.92229676,0.0011226982,0.00036340696,0.0003135708,0.0003842613,0.0033880975,0.007291571],"genre_scores_gemma":[0.5863514,0.0005980731,0.39300388,0.0010715986,0.00046145704,0.00038775805,0.00080779503,0.00037153272,0.01694648],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991391,0.00022138617,0.00005739213,0.00026517967,0.0002305862,0.00008628605],"domain_scores_gemma":[0.9962232,0.0017908994,0.00014456113,0.0005895824,0.0011360424,0.00011574676],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036356242,0.0012799255,0.001252582,0.0014738265,0.0005499361,0.0012150488,0.002944464,0.0023676192,0.0045648264],"category_scores_gemma":[0.00771891,0.00051345193,0.0013367528,0.0013587233,0.00067815627,0.0023976443,0.0020034725,0.0034735498,0.0012623732],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005423886,0.0006770456,0.0045634275,0.0001786487,0.00051431643,0.00018881484,0.00013334108,0.2822314,0.017553927,0.010115822,0.00975048,0.67355037],"study_design_scores_gemma":[0.00001789547,0.00009970381,0.0008566789,0.0000130476465,0.0000543328,0.000041567844,0.000010672247,0.99018914,0.0029056503,0.005289011,0.000509646,0.000012730214],"about_ca_topic_score_codex":0.005935259,"about_ca_topic_score_gemma":0.010375553,"teacher_disagreement_score":0.005935259,"about_ca_system_score_codex":0.0009143253,"about_ca_system_score_gemma":0.0011803467,"threshold_uncertainty_score":0.019227266},"labels":[],"label_agreement":null},{"id":"W4376279092","doi":"10.3389/fcomp.2023.1178040","title":"How social interactions can affect Modern Code Review","year":2023,"lang":"en","type":"article","venue":"Frontiers in Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Innopolis University; Russian Science Foundation","keywords":"Team software process; Process (computing); Computer science; Code review; Software quality; Quality (philosophy); Software; Personal software process; Software development; Software engineering; Software development process; Process management; Engineering; Software construction","score_opus":0.03196546461706035,"score_gpt":0.30753540873429414,"score_spread":0.2755699441172338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376279092","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9597773,0.001030192,0.0062332354,0.0058136173,0.00013975438,0.00010792793,0.00005876967,0.00009168823,0.026747555],"genre_scores_gemma":[0.9967056,0.000297213,0.0011415477,0.0004176205,0.00005723914,0.000050036164,0.000019899555,0.000025786298,0.0012850848],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9763668,0.017237438,0.0006302863,0.0012777321,0.0034847471,0.0010029987],"domain_scores_gemma":[0.9051148,0.06478529,0.017096544,0.0019018193,0.0074670403,0.0036343602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006291628,0.0004454999,0.0002932624,0.0016804425,0.0053426353,0.0048531964,0.0007273975,0.0015158535,0.003684277],"category_scores_gemma":[0.046478927,0.00036071433,0.00039159655,0.00083182065,0.0036628088,0.0028013985,0.0036369541,0.0014323838,0.00042354627],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032773928,0.0006275308,0.177035,0.0010451504,0.00026371313,0.002770362,0.6295081,0.0012303382,0.011005468,0.011143716,0.008062748,0.15698011],"study_design_scores_gemma":[0.000093739225,0.0010533421,0.3117584,0.0008781996,0.0004242037,0.002275628,0.50351834,0.006998089,0.0048857504,0.017124902,0.1505844,0.00040505038],"about_ca_topic_score_codex":0.0045174523,"about_ca_topic_score_gemma":0.006188229,"teacher_disagreement_score":0.006291628,"about_ca_system_score_codex":0.003344876,"about_ca_system_score_gemma":0.002778201,"threshold_uncertainty_score":0.033273697},"labels":[],"label_agreement":null},{"id":"W4376505229","doi":"10.1145/3597208","title":"An Empirical Study on GitHub Pull Requests’ Reactions","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; École de Technologie Supérieure","funders":"","keywords":"Computer science; Leverage (statistics); Source code; Set (abstract data type); Code review; Open source; Software; Empirical research; Code (set theory); Process (computing); Static program analysis; Software engineering; World Wide Web; Software development; Programming language; Artificial intelligence","score_opus":0.16806167547143225,"score_gpt":0.40876858222722956,"score_spread":0.2407069067557973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376505229","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9912498,0.00020528086,0.0028694742,0.0006376542,0.000033632037,0.00029955475,0.00029976454,0.000120596014,0.0042842817],"genre_scores_gemma":[0.9920844,0.00032808588,0.0033875396,0.0006443077,0.00005461203,0.000678525,0.0004967323,0.00012312306,0.0022027337],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9676914,0.018039312,0.0022112692,0.0023867965,0.008037245,0.0016340072],"domain_scores_gemma":[0.65195394,0.2472757,0.04651977,0.0085525485,0.039596867,0.0061010895],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.023799697,0.00075500726,0.0004951917,0.0038120197,0.0015008964,0.0035212093,0.0011883344,0.0017121118,0.0030114083],"category_scores_gemma":[0.15617275,0.00057331845,0.0004408941,0.002774784,0.0019558596,0.0038499467,0.002793472,0.0026892999,0.0016484093],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085761613,0.0016739968,0.51862043,0.0023496677,0.00010924943,0.0019553693,0.35882372,0.00053889805,0.011210864,0.0019949493,0.007592035,0.094273195],"study_design_scores_gemma":[0.00007704987,0.0014346876,0.649215,0.0011147591,0.00008664751,0.0012042723,0.29661113,0.0053727184,0.0069591794,0.0011755174,0.036508683,0.00024042116],"about_ca_topic_score_codex":0.0019550687,"about_ca_topic_score_gemma":0.002114036,"teacher_disagreement_score":0.996188,"about_ca_system_score_codex":0.0020232669,"about_ca_system_score_gemma":0.0017144049,"threshold_uncertainty_score":0.12586635},"labels":[],"label_agreement":null},{"id":"W4376606580","doi":"10.1109/saner56733.2023.00013","title":"Towards Understanding the Impacts of Textual Dissimilarity on Duplicate Bug Report Detection","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Software bug; Visualization; Construct (python library); Software regression; Security bug; Overhead (engineering); Domain (mathematical analysis); Software; Information retrieval; Data mining; Software development; Software quality","score_opus":0.07026126813592734,"score_gpt":0.3222827029850321,"score_spread":0.2520214348491048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376606580","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9157696,0.0026119358,0.07681176,0.00066379976,0.00007726952,0.00022663479,0.0009124015,0.0015094288,0.0014171498],"genre_scores_gemma":[0.9338094,0.00046276866,0.06318438,0.00014661517,0.000086156426,0.00010011949,0.0017179672,0.00017397935,0.0003185065],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9867857,0.005459252,0.0014879305,0.0024255484,0.0034355575,0.00040604043],"domain_scores_gemma":[0.80372447,0.13796574,0.027576352,0.013156098,0.015806263,0.001771022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011235038,0.00111483,0.0011807368,0.0074438513,0.00090381526,0.0029618698,0.0013874143,0.0014884652,0.00066144555],"category_scores_gemma":[0.14479496,0.0006613724,0.0009003476,0.0039076363,0.001347572,0.007443038,0.0023630878,0.0017973876,0.00038453724],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015207272,0.0010532342,0.5254158,0.0030059142,0.0007061848,0.0010586374,0.0073602074,0.018385306,0.047327593,0.0024384751,0.0049235015,0.38680437],"study_design_scores_gemma":[0.00024822354,0.0024628434,0.496647,0.0004963152,0.0010185398,0.0031507502,0.0062867673,0.41293538,0.053489577,0.011502806,0.011415139,0.00034676187],"about_ca_topic_score_codex":0.0030424902,"about_ca_topic_score_gemma":0.004013136,"teacher_disagreement_score":0.011235038,"about_ca_system_score_codex":0.00079311256,"about_ca_system_score_gemma":0.0009742513,"threshold_uncertainty_score":0.059417307},"labels":[],"label_agreement":null},{"id":"W4376606636","doi":"10.1109/saner56733.2023.00049","title":"A Multi-Step Learning Approach to Assist Code Review","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Snippet; Code review; Code (set theory); Source code; Process (computing); Task (project management); Artificial intelligence; Machine learning; Set (abstract data type); Software engineering; Static program analysis; Natural language processing; Data science; Programming language; Software development; Software","score_opus":0.07818152572049211,"score_gpt":0.3345728822930996,"score_spread":0.2563913565726075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376606636","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08259081,0.0025031788,0.8831053,0.0025934225,0.0005163173,0.0011321104,0.001246247,0.022400443,0.0039121853],"genre_scores_gemma":[0.4340997,0.0006232492,0.550007,0.0012089533,0.00031650453,0.0007610314,0.0029926964,0.0003543853,0.00963642],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9963871,0.0011230487,0.00026761042,0.0012423269,0.0008089876,0.0001708449],"domain_scores_gemma":[0.98452294,0.0074027576,0.0012260549,0.0010375488,0.005252293,0.0005583941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00393377,0.0018859745,0.0012936488,0.0031312334,0.00080430426,0.0014777388,0.0034082427,0.0022450641,0.0022182749],"category_scores_gemma":[0.017427528,0.0006697754,0.0013992204,0.0013584483,0.0005249537,0.0022651572,0.0022793836,0.0024493798,0.0023487315],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063739193,0.0011267215,0.009580111,0.0006557508,0.0002259579,0.0004917273,0.0006507616,0.060175065,0.015507197,0.00143195,0.01729811,0.89221936],"study_design_scores_gemma":[0.000075357224,0.00057900156,0.0018324796,0.00007682168,0.00012528322,0.00023076513,0.00013498538,0.9694276,0.014500165,0.0040837587,0.0088714985,0.000062270025],"about_ca_topic_score_codex":0.0035114132,"about_ca_topic_score_gemma":0.008175589,"teacher_disagreement_score":0.00393377,"about_ca_system_score_codex":0.0012366428,"about_ca_system_score_gemma":0.0027444009,"threshold_uncertainty_score":0.020804048},"labels":[],"label_agreement":null},{"id":"W4376606818","doi":"10.1109/saner56733.2023.00071","title":"Combining Contexts from Multiple Sources for Documentation-Specific Code Example Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates; University of Calgary","keywords":"Documentation; Internal documentation; Computer science; Code (set theory); Programming language; Compiler; Source code; Software documentation; Unit testing; Redundant code; Code generation; Software engineering; Software; Software development; Operating system; Key (lock); Set (abstract data type); Software development process; Software construction","score_opus":0.08059542835924342,"score_gpt":0.3064540948118102,"score_spread":0.22585866645256675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376606818","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39468738,0.0023999265,0.5566152,0.0013079832,0.0002980012,0.0011018589,0.0023006664,0.03162877,0.0096602645],"genre_scores_gemma":[0.52348626,0.00044351915,0.46493924,0.00038457656,0.00006149514,0.0005858695,0.0050446913,0.0017985145,0.0032558816],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99809784,0.0009376408,0.00011459419,0.00043949188,0.0003156928,0.000094753275],"domain_scores_gemma":[0.99103945,0.0054799537,0.00038276333,0.0018269684,0.0010633698,0.00020741747],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002241348,0.0012178923,0.00053894176,0.0016632505,0.00046898768,0.0010484782,0.0012591915,0.0011258953,0.0022801354],"category_scores_gemma":[0.018705774,0.0005461057,0.0008415795,0.0010164087,0.00048526638,0.0022650333,0.0023414753,0.001621685,0.0015033917],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007455731,0.00071262964,0.039141666,0.0016576635,0.00019073123,0.0014665839,0.0028221856,0.036991812,0.026780998,0.0051971837,0.01744072,0.86685234],"study_design_scores_gemma":[0.00025557054,0.0007522217,0.013930384,0.0007018074,0.00028138218,0.001747644,0.0012305056,0.84134173,0.058865815,0.022701038,0.05803912,0.00015276439],"about_ca_topic_score_codex":0.0014085383,"about_ca_topic_score_gemma":0.0048243776,"teacher_disagreement_score":0.0022801354,"about_ca_system_score_codex":0.000428916,"about_ca_system_score_gemma":0.0011509581,"threshold_uncertainty_score":0.011853516},"labels":[],"label_agreement":null},{"id":"W4376606868","doi":"10.1109/saner56733.2023.00077","title":"JTestMigBench and JTestMigTax: A benchmark and taxonomy for unit test migration","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Unit testing; Computer science; Reuse; Benchmark (surveying); Software engineering; Code (set theory); Taxonomy (biology); Code coverage; Test case; Code reuse; Scratch; Software; Software quality; Test (biology); Programming language; Machine learning; Software development; Engineering","score_opus":0.04112854843615728,"score_gpt":0.2681504518067819,"score_spread":0.2270219033706246,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376606868","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42426747,0.011548881,0.44308448,0.004059365,0.00083201623,0.004421914,0.014495331,0.057870798,0.03941984],"genre_scores_gemma":[0.35267335,0.0031673717,0.5869903,0.0010526081,0.00014044,0.0029772958,0.037000097,0.007895435,0.008103095],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98527896,0.0036100906,0.002559564,0.0016887172,0.0057920674,0.0010706146],"domain_scores_gemma":[0.946829,0.020387175,0.008146055,0.009553887,0.012682452,0.0024014025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009135823,0.0018434541,0.0007633312,0.009243351,0.0022706133,0.0028431166,0.003526745,0.0023034469,0.0019728313],"category_scores_gemma":[0.05855082,0.0010348911,0.0013789869,0.009549382,0.0017804278,0.004636799,0.0028551584,0.0027653016,0.0011433335],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016165397,0.001908092,0.122270145,0.0071441834,0.0003864966,0.0014834875,0.009383958,0.026117202,0.04123454,0.03600762,0.07355224,0.67889553],"study_design_scores_gemma":[0.000655498,0.005231391,0.18193105,0.005471949,0.0005517262,0.008598842,0.007811684,0.18040533,0.08594944,0.05391368,0.46864277,0.00083665823],"about_ca_topic_score_codex":0.009503453,"about_ca_topic_score_gemma":0.015377523,"teacher_disagreement_score":0.009503453,"about_ca_system_score_codex":0.0023722372,"about_ca_system_score_gemma":0.0044783414,"threshold_uncertainty_score":0.048315465},"labels":[],"label_agreement":null},{"id":"W4376632359","doi":"10.48550/arxiv.2305.07097","title":"Automated Smell Detection and Recommendation in Natural Language Requirements","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Usability; Ambiguity; Natural language; Quality (philosophy); Natural language processing; Requirements elicitation; Code smell; Domain (mathematical analysis); Recall; Requirements management; Requirements analysis; Non-functional testing; Artificial intelligence; Precision and recall; Natural (archaeology); Requirements engineering; Software engineering; Human–computer interaction; Software quality; Programming language; Linguistics; Software; Software development","score_opus":0.07626307219116117,"score_gpt":0.23863873021496715,"score_spread":0.162375658023806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376632359","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17074342,0.00064511644,0.77893317,0.00099631,0.00006135614,0.0010257877,0.0034806917,0.041020434,0.0030936843],"genre_scores_gemma":[0.28473556,0.0003038926,0.7053513,0.0003157555,0.000022546568,0.00041694273,0.006134427,0.0008221666,0.0018974062],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.989819,0.004028173,0.0011345878,0.001861913,0.0028702677,0.00028600797],"domain_scores_gemma":[0.94810206,0.03240288,0.007840171,0.0049735885,0.0061923615,0.00048894965],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046700714,0.0012961222,0.0010464664,0.0044413554,0.0005818634,0.0017543965,0.0016454988,0.0015498683,0.0018430075],"category_scores_gemma":[0.034796145,0.00094521686,0.0015009872,0.0016401941,0.0006143592,0.0022679328,0.0012799961,0.0011432839,0.0015555841],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007712948,0.0010236993,0.036947954,0.003956252,0.00031875656,0.0028645454,0.0032590108,0.036765933,0.15500467,0.004236357,0.021323498,0.73352796],"study_design_scores_gemma":[0.00017347804,0.00063434086,0.029499022,0.00052284513,0.00018507685,0.0024317894,0.0015068448,0.80341345,0.12637661,0.008007957,0.027017348,0.00023115287],"about_ca_topic_score_codex":0.003913121,"about_ca_topic_score_gemma":0.0082972115,"teacher_disagreement_score":0.0046700714,"about_ca_system_score_codex":0.0009962782,"about_ca_system_score_gemma":0.0016833496,"threshold_uncertainty_score":0.02469796},"labels":[],"label_agreement":null},{"id":"W4377100930","doi":"10.1016/j.jss.2023.111752","title":"Empirical analysis of security-related code reviews in npm packages","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Concordia University; University of Waterloo","funders":"","keywords":"Computer science; Code review; Software security assurance; Identification (biology); Code (set theory); Relation (database); Quality (philosophy); Domain (mathematical analysis); Risk analysis (engineering); Software engineering; Software; Computer security; Static program analysis; Software development; Security service; Information security; Database; Business; Programming language","score_opus":0.03924305586657146,"score_gpt":0.3280759501969659,"score_spread":0.28883289433039444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4377100930","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99702007,0.00037280098,0.00036092126,0.00027055867,0.000013276962,0.000049608247,0.00031452804,0.000021764115,0.0015763626],"genre_scores_gemma":[0.9985222,0.00011883225,0.00028560878,0.00007699985,0.000016381022,0.00004578919,0.000280415,0.000020961299,0.0006328261],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9844828,0.006554992,0.0012963953,0.0015301171,0.00525953,0.000876162],"domain_scores_gemma":[0.43080685,0.37972525,0.12887554,0.00933326,0.044800557,0.006458571],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0115685,0.0002725441,0.00035626994,0.005039479,0.0009296533,0.0015811597,0.0012790626,0.0013259711,0.0041063065],"category_scores_gemma":[0.25978017,0.00035675036,0.0005522191,0.0033304857,0.0013019391,0.0021101024,0.0014268707,0.002003402,0.0008655071],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046520043,0.00046289983,0.97301435,0.00029007203,0.0001426914,0.00022781828,0.003550278,0.00051669934,0.00061054836,0.00065273675,0.001974815,0.018091982],"study_design_scores_gemma":[0.000020934895,0.0003246995,0.99000484,0.00016735615,0.00010955228,0.00035924357,0.003562481,0.0023393396,0.0006263873,0.00024014422,0.0022200942,0.000024880539],"about_ca_topic_score_codex":0.0077691106,"about_ca_topic_score_gemma":0.012793421,"teacher_disagreement_score":0.9884315,"about_ca_system_score_codex":0.0020425927,"about_ca_system_score_gemma":0.0028871407,"threshold_uncertainty_score":0.06118083},"labels":[],"label_agreement":null},{"id":"W4377103668","doi":"10.1145/3592623","title":"Asm2Seq: Explainable Assembly Code Functional Summary Generation for Reverse Engineering and Vulnerability Analysis","year":2023,"lang":"en","type":"article","venue":"Digital Threats Research and Practice","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; Queen's University","funders":"","keywords":"Computer science; Reverse engineering; Firmware; Source code; Assembly language; Automatic summarization; Vulnerability (computing); Code review; Code (set theory); Static program analysis; Leverage (statistics); Software; Process (computing); Context (archaeology); Software engineering; Artificial intelligence; Programming language; Software development; Computer security; Operating system","score_opus":0.1742418370316224,"score_gpt":0.39432737323836237,"score_spread":0.22008553620673996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4377103668","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04095696,0.0010890163,0.72778153,0.00070279435,0.00037745084,0.0007320956,0.030566795,0.19376093,0.0040325057],"genre_scores_gemma":[0.13888796,0.0003826679,0.75417036,0.00040464295,0.00013523415,0.00088201667,0.09349472,0.0058724363,0.0057699685],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99890804,0.00025393878,0.00009086213,0.0003868913,0.00028379727,0.00007646646],"domain_scores_gemma":[0.9966272,0.0014122095,0.00026985767,0.0008131506,0.00078045635,0.0000972435],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011648529,0.002462589,0.0006603137,0.0029350522,0.0005642287,0.0009635,0.0018095183,0.0013554139,0.009511383],"category_scores_gemma":[0.0075865514,0.0004981758,0.0013322174,0.0013022132,0.0004584151,0.001707744,0.0016486022,0.0014077794,0.005161437],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007554044,0.00035798308,0.007751052,0.0014468916,0.0002066531,0.00093774084,0.00082653674,0.02903477,0.041298334,0.006006016,0.15153341,0.7598452],"study_design_scores_gemma":[0.00025559665,0.0005678038,0.007251113,0.00013989976,0.00017384574,0.0010019896,0.00040956054,0.77182007,0.08632068,0.020591069,0.1113254,0.0001429465],"about_ca_topic_score_codex":0.004723047,"about_ca_topic_score_gemma":0.010287296,"teacher_disagreement_score":0.009511383,"about_ca_system_score_codex":0.0007660277,"about_ca_system_score_gemma":0.0014552369,"threshold_uncertainty_score":0.031818748},"labels":[],"label_agreement":null},{"id":"W4377139078","doi":"10.1007/s10664-023-10300-3","title":"Learning to Predict Code Review Completion Time In Modern Code Review","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code review; Computer science; Context (archaeology); Coding (social sciences); Code (set theory); Process (computing); Best practice; Quality (philosophy); Software engineering; Software quality; Software development; Software","score_opus":0.03732351305167277,"score_gpt":0.3175656528412004,"score_spread":0.28024213978952767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4377139078","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97085774,0.0005015635,0.024977919,0.0004826365,0.00007640404,0.00009121067,0.00073908555,0.00081076537,0.0014627957],"genre_scores_gemma":[0.9864433,0.000109528664,0.01052204,0.00007837543,0.00005275942,0.000058031488,0.001357423,0.000048826045,0.0013296094],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9984358,0.0005547443,0.00014111439,0.00040609972,0.00031867958,0.00014356249],"domain_scores_gemma":[0.9263188,0.05645147,0.0065001417,0.002237249,0.006292188,0.0022002396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038634534,0.00057129253,0.00044494582,0.0019060584,0.00035848533,0.0009150347,0.00061305193,0.00091356545,0.0022031676],"category_scores_gemma":[0.056065157,0.00026605165,0.0005608508,0.0008521314,0.00025194953,0.0014368958,0.00062250666,0.0016975058,0.0011637665],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015603822,0.0022702147,0.6518449,0.00023432364,0.0002664933,0.00011012325,0.00028702433,0.055657864,0.0027719259,0.0007650135,0.01012071,0.274111],"study_design_scores_gemma":[0.00010830449,0.0010663455,0.13425016,0.00006630853,0.00011230194,0.00014701484,0.00018673335,0.85565615,0.0035169055,0.0032235596,0.0016208882,0.00004542145],"about_ca_topic_score_codex":0.00588287,"about_ca_topic_score_gemma":0.010571091,"teacher_disagreement_score":0.00588287,"about_ca_system_score_codex":0.0008089977,"about_ca_system_score_gemma":0.0015613656,"threshold_uncertainty_score":0.020432115},"labels":[],"label_agreement":null},{"id":"W4377194214","doi":"10.1142/s2811032323500017","title":"A Survey on Code Representation","year":2023,"lang":"en","type":"article","venue":"World Scientific Annual Review of Artificial Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Embedding; Computer science; Benchmark (surveying); Representation (politics); Artificial intelligence; Code (set theory); Feature learning; Machine learning; Natural language processing; Space (punctuation); Programming language","score_opus":0.1135828696490033,"score_gpt":0.39746013919591844,"score_spread":0.28387726954691517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4377194214","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01080364,0.369234,0.54125273,0.010543093,0.001769996,0.00031462678,0.0034235618,0.0044814604,0.058176786],"genre_scores_gemma":[0.07015269,0.5534887,0.32929063,0.0035343564,0.0025280884,0.0006018519,0.0142636355,0.0016774596,0.024462529],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979091,0.0004216528,0.00020492874,0.00039918834,0.00093012897,0.00013483486],"domain_scores_gemma":[0.99651355,0.0019154061,0.00019591668,0.00052879396,0.0007485541,0.0000978751],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014476031,0.0011310455,0.0010057747,0.0053245504,0.000707131,0.0029127684,0.0022216465,0.0014123349,0.009771743],"category_scores_gemma":[0.010057304,0.00061919243,0.001243974,0.009927519,0.001132137,0.007833354,0.0020169867,0.0021909694,0.0044418024],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035885994,0.000072388415,0.0012292219,0.0021015604,0.00002769706,0.00006890495,0.00019232319,0.0055413195,0.0008661723,0.086029485,0.04138457,0.8624505],"study_design_scores_gemma":[0.000015614323,0.000101596546,0.0016457598,0.0026521147,0.000048090045,0.00097298244,0.00036203692,0.03479853,0.0028994128,0.15996301,0.79646397,0.00007687215],"about_ca_topic_score_codex":0.0040059877,"about_ca_topic_score_gemma":0.0026496034,"teacher_disagreement_score":0.009771743,"about_ca_system_score_codex":0.0013831872,"about_ca_system_score_gemma":0.0029513778,"threshold_uncertainty_score":0.03268969},"labels":[],"label_agreement":null},{"id":"W4378072131","doi":"10.1007/s10664-023-10287-x","title":"Rubbing salt in the wound? A large-scale investigation into the effects of refactoring on security","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Ministero dell’Istruzione, dell’Università e della Ricerca; Università degli Studi di Salerno; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; National Science Foundation","keywords":"Code refactoring; Computer science; Technical debt; Source code; Software engineering; Software; Code (set theory); Source lines of code; Software development; Computer security; Programming language; Set (abstract data type)","score_opus":0.017211188368017438,"score_gpt":0.2799320263002199,"score_spread":0.26272083793220247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378072131","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99636227,0.00029024578,0.002251615,0.00027835625,0.000008160935,0.000054812957,0.00015603905,0.000049283153,0.0005490627],"genre_scores_gemma":[0.99580646,0.00025210134,0.0032043816,0.00013501362,0.00001387388,0.00004261137,0.00024373429,0.00003711621,0.0002645887],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98595726,0.007466434,0.000960359,0.0017568826,0.0033615925,0.00049741706],"domain_scores_gemma":[0.62582564,0.29678512,0.037133045,0.021596598,0.01689694,0.0017627162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0164219,0.0004239197,0.0004021128,0.0028080093,0.0008084961,0.0013913534,0.0010741246,0.0006583554,0.0014700937],"category_scores_gemma":[0.12210085,0.00030998778,0.00074050034,0.0029736722,0.0016843796,0.0029653993,0.0016046123,0.0013287719,0.0004252694],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058758416,0.002518467,0.8291266,0.00067656,0.00043219904,0.0006063558,0.007006134,0.0028860162,0.0049778833,0.0013299008,0.0021847792,0.14766756],"study_design_scores_gemma":[0.000060522012,0.0019230022,0.9670313,0.00038794678,0.00024499965,0.0004213318,0.0074517685,0.011387828,0.006628375,0.0012721225,0.0031261311,0.00006462959],"about_ca_topic_score_codex":0.0034487941,"about_ca_topic_score_gemma":0.0046633584,"teacher_disagreement_score":0.0164219,"about_ca_system_score_codex":0.0010560267,"about_ca_system_score_gemma":0.0012795074,"threshold_uncertainty_score":0.08684832},"labels":[],"label_agreement":null},{"id":"W4378373692","doi":"10.1109/icst57152.2023.00018","title":"Embedding Context as Code Dependencies for Neural Program Repair","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; ENCODE; Embedding; Source code; Code (set theory); Graph; Dependency graph; Artificial intelligence; Programming language; Theoretical computer science; Representation (politics); Natural language processing","score_opus":0.048594963281501695,"score_gpt":0.3578197693393004,"score_spread":0.3092248060577987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378373692","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.141648,0.0018172077,0.8426518,0.0007359171,0.00012826791,0.00007411332,0.000805675,0.01031655,0.0018223429],"genre_scores_gemma":[0.8164232,0.00059897266,0.17690548,0.0002779455,0.00006682975,0.00010830701,0.0020984244,0.0005311125,0.0029897569],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996468,0.000064928194,0.000018127199,0.00015462082,0.000070683374,0.000044907883],"domain_scores_gemma":[0.998922,0.0005125563,0.00016900823,0.00019578393,0.00014955632,0.000051196177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00042934166,0.0011691244,0.00055225857,0.0013625273,0.00037450416,0.00041820883,0.0015286839,0.0010471093,0.0018690972],"category_scores_gemma":[0.0033225068,0.00048513827,0.0007487137,0.00087562663,0.0006582688,0.0019708965,0.0008842402,0.0019236066,0.00045254742],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021942859,0.00018907936,0.0036871908,0.00022825888,0.0000738347,0.0001628869,0.00013876152,0.5887717,0.012273192,0.007989535,0.004607982,0.3816581],"study_design_scores_gemma":[0.0000073576293,0.00004000455,0.00037496723,0.00001153006,0.000016612346,0.000022545457,0.00001455565,0.98645127,0.002642889,0.009624777,0.00078604114,0.000007463495],"about_ca_topic_score_codex":0.0089181075,"about_ca_topic_score_gemma":0.01916161,"teacher_disagreement_score":0.0089181075,"about_ca_system_score_codex":0.0010512418,"about_ca_system_score_gemma":0.001073674,"threshold_uncertainty_score":0.017732382},"labels":[],"label_agreement":null},{"id":"W4378474030","doi":"10.1007/s10664-023-10317-8","title":"Using the uniqueness of global identifiers to determine the provenance of Python software source code","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Identifier; Python (programming language); Computer science; Source code; Software; Open source; Source lines of code; Database; Programming language; Operating system","score_opus":0.05727007110122434,"score_gpt":0.33277357167196664,"score_spread":0.2755035005707423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378474030","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.66259,0.00045204573,0.32322943,0.00062896113,0.00017198604,0.00019985826,0.0020370928,0.0011579938,0.009532563],"genre_scores_gemma":[0.9309652,0.00015145939,0.06631658,0.00006989756,0.0000450957,0.0000984898,0.0010390227,0.0002521086,0.001062152],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98975414,0.0034869919,0.0010934989,0.0018070904,0.0032627871,0.0005953431],"domain_scores_gemma":[0.8834149,0.05284267,0.017014861,0.02336885,0.021308923,0.002049863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013666544,0.00032536977,0.00054424565,0.0048348396,0.0017105783,0.0032360416,0.0008765295,0.0010549532,0.0017588424],"category_scores_gemma":[0.1471189,0.0005533031,0.00058287865,0.0037840733,0.0023646825,0.009539221,0.004887024,0.0021877591,0.00060321327],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083066087,0.000186664,0.6576696,0.0005456737,0.00023421776,0.0005899842,0.011130986,0.0092903925,0.016691424,0.11319984,0.0034025926,0.18622808],"study_design_scores_gemma":[0.0001973177,0.0006640716,0.28691274,0.000862926,0.0004738396,0.002018934,0.009961264,0.22092894,0.076261185,0.35676184,0.044577908,0.0003790948],"about_ca_topic_score_codex":0.0045536123,"about_ca_topic_score_gemma":0.0054886835,"teacher_disagreement_score":0.013666544,"about_ca_system_score_codex":0.0010733886,"about_ca_system_score_gemma":0.003994755,"threshold_uncertainty_score":0.07227641},"labels":[],"label_agreement":null},{"id":"W4378770585","doi":"10.48550/arxiv.2305.17286","title":"A Study of Documentation for Software Architecture","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Software documentation; Computer science; Internal documentation; Context (archaeology); World Wide Web; Architecture; Software; Software engineering; Software system; Information retrieval; Programming language; Software construction","score_opus":0.10322875815007611,"score_gpt":0.2429092687851422,"score_spread":0.1396805106350661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378770585","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98978525,0.00046027414,0.0040679355,0.00082529214,0.00000912342,0.00004023131,0.000014802813,0.000011978267,0.004785092],"genre_scores_gemma":[0.99390054,0.00040743887,0.004236577,0.00022457863,0.000019259996,0.00007200191,0.00002871888,0.000011714375,0.001099192],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9896365,0.007983162,0.0003867893,0.00045678567,0.0013259598,0.00021082078],"domain_scores_gemma":[0.78348166,0.18058854,0.022220057,0.004455354,0.006155115,0.0030992778],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011689699,0.00029633884,0.00040708095,0.0011564608,0.0018906002,0.0027469806,0.00056376075,0.0010098505,0.0018106418],"category_scores_gemma":[0.092344366,0.0003511193,0.00026415146,0.0016279753,0.0028434235,0.003867329,0.0013864092,0.0020515537,0.00024256195],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024630685,0.0021752762,0.18889898,0.00080315425,0.000060084814,0.0011002468,0.65277654,0.00080668536,0.0057334807,0.026993185,0.0017807969,0.11862536],"study_design_scores_gemma":[0.00025587506,0.0059205154,0.46510413,0.0019087105,0.0001263504,0.0042801225,0.36785388,0.010740045,0.0063682836,0.04424338,0.092996344,0.00020240572],"about_ca_topic_score_codex":0.0011047454,"about_ca_topic_score_gemma":0.0014591723,"teacher_disagreement_score":0.011689699,"about_ca_system_score_codex":0.0014615483,"about_ca_system_score_gemma":0.0020677329,"threshold_uncertainty_score":0.06182176},"labels":[],"label_agreement":null},{"id":"W4378977640","doi":"10.1016/j.jss.2023.111767","title":"BPEL process defects prediction using multi-objective evolutionary search","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Business Process Execution Language; Computer science; Artifact (error); Process (computing); Software engineering; Web service; Business process; Orchestration; Software; Software portability; Genetic programming; Symbolic regression; Machine learning; Artificial intelligence; Data mining; Service-oriented architecture; Programming language; Work in process; Engineering","score_opus":0.038743592987931175,"score_gpt":0.30080741990111104,"score_spread":0.26206382691317986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378977640","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47839937,0.0006966193,0.51677126,0.00018896592,0.00006107173,0.000107801876,0.00017070802,0.0012030544,0.002401169],"genre_scores_gemma":[0.93094295,0.000110050416,0.06729262,0.00003737513,0.000014356682,0.00006665697,0.0002089171,0.000048283546,0.0012787385],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99936396,0.00014181246,0.00004864086,0.0001577266,0.00020551756,0.00008241296],"domain_scores_gemma":[0.9977875,0.0013672488,0.0002177982,0.000082948405,0.00047640235,0.00006811291],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014176226,0.0011516601,0.0013998442,0.0025566234,0.0004458389,0.0008549152,0.0011399275,0.0014610841,0.0011502272],"category_scores_gemma":[0.0034155776,0.0004998909,0.0010438821,0.0011018467,0.00027362772,0.00095794536,0.0005653428,0.00073340477,0.0002079489],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019196811,0.00026319353,0.005062751,0.00006866508,0.00010563439,0.0000984249,0.000023741697,0.91318494,0.002310548,0.00046512968,0.00029219067,0.07793279],"study_design_scores_gemma":[0.000004691997,0.000019603347,0.00019407076,0.0000019624065,0.000007531552,0.000005628442,0.00000213784,0.99944633,0.00023527954,0.000067303314,0.000013984411,0.0000013709854],"about_ca_topic_score_codex":0.00769434,"about_ca_topic_score_gemma":0.004492507,"teacher_disagreement_score":0.00769434,"about_ca_system_score_codex":0.00066289894,"about_ca_system_score_gemma":0.0009289392,"threshold_uncertainty_score":0.015299082},"labels":[],"label_agreement":null},{"id":"W4379014622","doi":"10.1145/3603110","title":"Dependency Update Strategies and Package Characteristics","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Concordia University","funders":"","keywords":"Computer science; Dependency (UML); Software versioning; Software engineering; Dilemma; Software; Programming language","score_opus":0.07638883686026998,"score_gpt":0.32458871075136997,"score_spread":0.2481998738911,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379014622","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98067284,0.00021998509,0.014158143,0.00019947423,0.000014565683,0.00008283663,0.00052968215,0.000340938,0.00378138],"genre_scores_gemma":[0.99187696,0.00008012418,0.0061546317,0.000027342243,0.000007102756,0.000042986067,0.00068527507,0.0001331982,0.0009924062],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99552745,0.0015024488,0.00035357068,0.00081350066,0.0014810555,0.00032198933],"domain_scores_gemma":[0.8791574,0.081227206,0.018567218,0.008705438,0.010157777,0.0021849296],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061071836,0.0005479517,0.00034075064,0.002146082,0.0006246971,0.001935728,0.0009451137,0.0005555774,0.0022877129],"category_scores_gemma":[0.07588558,0.0005679655,0.0004859971,0.002056025,0.00078942836,0.0039345887,0.0012484863,0.0009952011,0.00072505034],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014909216,0.00016686115,0.91723734,0.00010772584,0.00009707988,0.00021786802,0.0018064934,0.010751565,0.0019489719,0.0016038961,0.001970578,0.06394254],"study_design_scores_gemma":[0.000022863383,0.0003578591,0.9054407,0.00007349465,0.00016406797,0.0008961281,0.00211362,0.0743191,0.0040715844,0.0035100945,0.008936467,0.000094059826],"about_ca_topic_score_codex":0.0050614844,"about_ca_topic_score_gemma":0.008134998,"teacher_disagreement_score":0.0061071836,"about_ca_system_score_codex":0.0010712668,"about_ca_system_score_gemma":0.0007990565,"threshold_uncertainty_score":0.032298267},"labels":[],"label_agreement":null},{"id":"W4379115941","doi":"10.23919/date56975.2023.10137086","title":"Benchmarking Large Language Models for Automated Verilog RTL Code Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":142,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Army Research Office; National Science Foundation","keywords":"Computer science; Verilog; Programming language; Compiler; Correctness; Construct (python library); Scripting language; Embedded system; Field-programmable gate array","score_opus":0.04206808642623612,"score_gpt":0.31821297014393424,"score_spread":0.27614488371769813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379115941","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34022298,0.004060035,0.37689254,0.0025853328,0.0012227905,0.0014792252,0.030006759,0.22532383,0.018206539],"genre_scores_gemma":[0.4897899,0.0011669964,0.39918396,0.001266178,0.000105621,0.0012700207,0.08678579,0.014481896,0.005949599],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99573034,0.001530465,0.0004106837,0.00090458133,0.001088966,0.000334906],"domain_scores_gemma":[0.98335016,0.009503154,0.0007678891,0.0035172557,0.0025300826,0.0003314691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004708408,0.0021168573,0.00055932323,0.0016699114,0.0005733717,0.0016399068,0.0034796807,0.0018474382,0.007289788],"category_scores_gemma":[0.025979135,0.0011359162,0.0016630319,0.0012736025,0.0011384533,0.0034561977,0.0019828414,0.0030381181,0.0049671987],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013844053,0.0016844542,0.0206165,0.0035359836,0.00046475857,0.00054139487,0.0007771144,0.40301317,0.029414637,0.0137732355,0.1505856,0.37420878],"study_design_scores_gemma":[0.00029218377,0.00044047972,0.002546144,0.00022660056,0.00008273758,0.00021154352,0.00018539432,0.9259245,0.032482132,0.00840903,0.029127399,0.00007190147],"about_ca_topic_score_codex":0.007382188,"about_ca_topic_score_gemma":0.014732745,"teacher_disagreement_score":0.007382188,"about_ca_system_score_codex":0.0021302598,"about_ca_system_score_gemma":0.0036004146,"threshold_uncertainty_score":0.024900734},"labels":[],"label_agreement":null},{"id":"W4379391205","doi":"10.5121/csit.2023.130810","title":"Analyzing Emotional Contagion in Commit Messages of Open-Source Software Repositories","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Commit; Emotional contagion; Software; Affect (linguistics); Productivity; Computer science; Open source software; Creativity; Software development; Time zone; Process (computing); Psychology; Social psychology; Communication","score_opus":0.027046312835567243,"score_gpt":0.30135844593399963,"score_spread":0.2743121330984324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379391205","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99649847,0.00009821358,0.0013053817,0.00014488571,0.000027380533,0.000042906784,0.00009144033,0.00004667982,0.0017446101],"genre_scores_gemma":[0.9971354,0.00008139762,0.0013119007,0.00011270921,0.000039689592,0.000095710464,0.00019694817,0.00003288711,0.0009932767],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99762803,0.0011337464,0.00014008746,0.00023117122,0.000657643,0.00020931379],"domain_scores_gemma":[0.96630114,0.024680143,0.004553151,0.00091486186,0.0025359422,0.0010147169],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020076088,0.0003375603,0.00035409356,0.0014281044,0.00084930065,0.0016460991,0.0004149867,0.0008729479,0.0018497646],"category_scores_gemma":[0.03384518,0.00017564808,0.00022330179,0.0008795982,0.00054486125,0.0013375602,0.0019171331,0.0012760037,0.00041659814],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003007245,0.0013446102,0.5539575,0.0012819732,0.0003055411,0.0019437632,0.20204331,0.0020526284,0.04165419,0.0029367672,0.0051937867,0.18427868],"study_design_scores_gemma":[0.000044546057,0.00055677927,0.9292081,0.00020145853,0.0001364726,0.0004961178,0.0469822,0.01043697,0.0037741198,0.002095261,0.0059445887,0.00012333694],"about_ca_topic_score_codex":0.0015116883,"about_ca_topic_score_gemma":0.001984078,"teacher_disagreement_score":0.0020076088,"about_ca_system_score_codex":0.00055028737,"about_ca_system_score_gemma":0.0003240445,"threshold_uncertainty_score":0.010617375},"labels":[],"label_agreement":null},{"id":"W4379930888","doi":"10.1109/iciccs56967.2023.10142534","title":"Explainable Software Defect Prediction from Cross Company Project Metrics using Machine Learning","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Fanshawe College","funders":"","keywords":"Computer science; Software bug; Schedule; Machine learning; Sizing; Software; Predictive modelling; Product metric; Transparency (behavior); Software metric; Artificial intelligence; Class (philosophy); Data mining; Software engineering; Software development; Software quality","score_opus":0.057804348197871336,"score_gpt":0.3175388171236185,"score_spread":0.25973446892574714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379930888","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55957395,0.000420096,0.4359815,0.00047952047,0.000026933429,0.0000831425,0.0011329976,0.0008422246,0.0014596117],"genre_scores_gemma":[0.96954745,0.00010778774,0.028836511,0.000019828383,0.000014127515,0.00004400668,0.001020455,0.000022713331,0.00038717885],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992066,0.0003273492,0.00005205061,0.0001876874,0.00016121846,0.00006515665],"domain_scores_gemma":[0.9850797,0.01162752,0.0012844478,0.0009806005,0.00087134744,0.00015632868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023645316,0.00077033136,0.00042704708,0.0023349016,0.00018705275,0.0007933616,0.0007759646,0.0006380748,0.0009324156],"category_scores_gemma":[0.012030564,0.00022765661,0.0006712379,0.0013393271,0.0002554834,0.0014182349,0.00063277525,0.00091066316,0.00014067274],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013102182,0.00024779284,0.08281898,0.00011910951,0.00020698278,0.00018734644,0.00020500725,0.7739961,0.0008548318,0.0077758394,0.0013708387,0.13208614],"study_design_scores_gemma":[0.0000045329075,0.00003166413,0.0046729553,0.000007850446,0.000013351944,0.000016330401,0.000017476124,0.98971725,0.00027633135,0.005062675,0.00017400914,0.000005522773],"about_ca_topic_score_codex":0.0038581975,"about_ca_topic_score_gemma":0.0069366274,"teacher_disagreement_score":0.0038581975,"about_ca_system_score_codex":0.00082417653,"about_ca_system_score_gemma":0.0007948652,"threshold_uncertainty_score":0.012504995},"labels":[],"label_agreement":null},{"id":"W4380048657","doi":"10.1007/s10472-023-09844-3","title":"Quantifying the relationship between software design principles and performance in Jason: a case study with simulated mobile robots","year":2023,"lang":"en","type":"article","venue":"Annals of Mathematics and Artificial Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Maintainability; Cyclomatic complexity; Cohesion (chemistry); Computer science; Code (set theory); Separation of concerns; Coupling (piping); Software; Software engineering; Artificial intelligence; Reliability engineering; Programming language; Engineering","score_opus":0.4450014916976675,"score_gpt":0.4113975001879414,"score_spread":0.033603991509726094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380048657","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99738425,0.000013152562,0.0013490173,0.000049343493,0.0000019465003,0.00002959273,0.000015163532,0.000016532256,0.0011409553],"genre_scores_gemma":[0.9950294,0.000020060128,0.004226461,0.000007875846,0.0000013855275,0.00003943387,0.000035736975,0.000011314392,0.0006283664],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99857473,0.0007869708,0.000060686638,0.00013965802,0.00028240902,0.00015552189],"domain_scores_gemma":[0.976473,0.01920835,0.0014351537,0.0007438102,0.0013376944,0.00080194446],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035638108,0.0005670443,0.0003582957,0.0009501495,0.0011801707,0.0018442834,0.0010355632,0.0014283621,0.0016248178],"category_scores_gemma":[0.01743629,0.00034446002,0.000293055,0.00079113105,0.0016927449,0.0017908957,0.0012574352,0.0008679922,0.00027293962],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0058869906,0.009951797,0.23153494,0.0008999109,0.00029274484,0.0026999542,0.023394762,0.5536707,0.03196195,0.024049863,0.0023535062,0.11330292],"study_design_scores_gemma":[0.0004777112,0.011995591,0.13365184,0.00011067587,0.00023910914,0.00044930255,0.020356081,0.79551834,0.02061968,0.01055467,0.0057919794,0.00023502391],"about_ca_topic_score_codex":0.005178253,"about_ca_topic_score_gemma":0.010567692,"teacher_disagreement_score":0.005178253,"about_ca_system_score_codex":0.0019289053,"about_ca_system_score_gemma":0.00141462,"threshold_uncertainty_score":0.018847406},"labels":[],"label_agreement":null},{"id":"W4380520359","doi":"10.1109/tse.2023.3285743","title":"STRE: An Automated Approach to Suggesting App Developers When to Stop Reading Reviews","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Fundo para o Desenvolvimento das Ciências e da Tecnologia; China Postdoctoral Science Foundation","keywords":"Computer science; Reading (process); Upload; Categorization; World Wide Web; Complement (music); Data science; Artificial intelligence","score_opus":0.039493449767354936,"score_gpt":0.2921104368037568,"score_spread":0.2526169870364019,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380520359","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14010514,0.0077937907,0.50420415,0.003042444,0.0018219817,0.005065976,0.019341573,0.3025094,0.016115516],"genre_scores_gemma":[0.25184172,0.001236503,0.70723677,0.0010152603,0.0005985876,0.0012413935,0.015792957,0.0026699118,0.018366931],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941362,0.0013485475,0.0007134914,0.0016688505,0.0019033443,0.00022948305],"domain_scores_gemma":[0.9691749,0.012210709,0.005215107,0.0019490274,0.010370915,0.0010792708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034935377,0.0034186451,0.0017693723,0.009945918,0.0010390042,0.0022635176,0.0026627306,0.001780393,0.004945308],"category_scores_gemma":[0.024979094,0.0009846794,0.0009640992,0.0030048115,0.00044285486,0.0027493904,0.0015379209,0.0015819488,0.0076220613],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009520223,0.0007526611,0.021572609,0.0020891505,0.00025556693,0.00074600155,0.001987172,0.003665518,0.023949549,0.0010491249,0.11443197,0.8285486],"study_design_scores_gemma":[0.00087763363,0.0029002998,0.063588046,0.000899997,0.0011068981,0.0031590245,0.003284667,0.5872862,0.06736445,0.00677787,0.26188198,0.0008730623],"about_ca_topic_score_codex":0.009737097,"about_ca_topic_score_gemma":0.030253671,"teacher_disagreement_score":0.009945918,"about_ca_system_score_codex":0.0009383641,"about_ca_system_score_gemma":0.003422146,"threshold_uncertainty_score":0.01936084},"labels":[],"label_agreement":null},{"id":"W4380568713","doi":"10.1145/3593434.3593475","title":"Developers’ Perception of GitHub Actions: A Survey Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Documentation; Computer science; Debugging; Usability; Workflow; Action (physics); Process (computing); Visibility; World Wide Web; Perception; Software engineering; Data science; Human–computer interaction; Database; Programming language","score_opus":0.08776789958189843,"score_gpt":0.33645096473283564,"score_spread":0.24868306515093722,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380568713","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99502844,0.00065889617,0.0011864882,0.0005552328,0.000013599509,0.00006083524,0.00035261342,0.00006635438,0.0020775653],"genre_scores_gemma":[0.99602604,0.0007246916,0.0014160904,0.00027564328,0.000014201267,0.00012194373,0.00044131768,0.000041926156,0.00093811133],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9935696,0.0027369764,0.00069363485,0.00054259854,0.0019285899,0.0005286259],"domain_scores_gemma":[0.94946975,0.027832562,0.009763521,0.0012399241,0.009120872,0.0025734466],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0074776253,0.00026392564,0.0003614181,0.0023812186,0.00062673184,0.0012072755,0.0004813225,0.0006008682,0.0010832249],"category_scores_gemma":[0.035499457,0.00030307684,0.00038475374,0.0015933387,0.00070498994,0.0016202758,0.0013997402,0.00074343145,0.00036954402],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015036334,0.0001236579,0.85160047,0.00064610364,0.00006743539,0.0005952453,0.08001987,0.00019055714,0.0017332664,0.0003450554,0.0035454512,0.060982518],"study_design_scores_gemma":[0.000020581709,0.0003397733,0.85329175,0.00049406325,0.000073761556,0.001149878,0.11551462,0.001541421,0.00096660043,0.00029383934,0.026223512,0.00009025706],"about_ca_topic_score_codex":0.0061324863,"about_ca_topic_score_gemma":0.00861209,"teacher_disagreement_score":0.99252236,"about_ca_system_score_codex":0.0010611189,"about_ca_system_score_gemma":0.001191352,"threshold_uncertainty_score":0.039545953},"labels":[],"label_agreement":null},{"id":"W4380877364","doi":"10.1007/s10270-023-01113-5","title":"Repository mining for changes in Simulink and Stateflow models","year":2023,"lang":"en","type":"article","venue":"Software & Systems Modeling","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Stateflow; Computer science; Automotive industry; MATLAB; Software; Model-based design; Software engineering; Simulation; Programming language; Engineering","score_opus":0.05202037005972848,"score_gpt":0.28110912341442795,"score_spread":0.22908875335469947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380877364","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43947384,0.004065387,0.35349965,0.0018266077,0.00095555274,0.0011887884,0.103290424,0.083153106,0.012546673],"genre_scores_gemma":[0.65888566,0.00157724,0.2029977,0.00022502219,0.00010572993,0.0005096596,0.12703204,0.0032054242,0.00546159],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99412256,0.0005762597,0.00083901203,0.0013839542,0.0027592648,0.00031892353],"domain_scores_gemma":[0.9767595,0.008201026,0.0035041883,0.0067603383,0.0041679274,0.0006070178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034103098,0.0013466303,0.0011306976,0.012912205,0.001055172,0.0028028789,0.0031416842,0.0014144294,0.0034084436],"category_scores_gemma":[0.03235663,0.00076109334,0.0028605477,0.006439176,0.0006400888,0.00500882,0.002414547,0.0021278833,0.001637567],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012582162,0.0010817804,0.18854992,0.003950484,0.0011325463,0.0046471036,0.002756811,0.067021064,0.020231605,0.017975062,0.059458587,0.6319368],"study_design_scores_gemma":[0.0002013095,0.00068399194,0.064114586,0.0013680941,0.0017832262,0.0036927483,0.0020062465,0.68927616,0.082535066,0.021190539,0.1328316,0.0003164649],"about_ca_topic_score_codex":0.012640722,"about_ca_topic_score_gemma":0.02190328,"teacher_disagreement_score":0.012912205,"about_ca_system_score_codex":0.0014981776,"about_ca_system_score_gemma":0.0036303492,"threshold_uncertainty_score":0.025134265},"labels":[],"label_agreement":null},{"id":"W4381185676","doi":"10.1017/pds.2023.107","title":"CONNECTING DESIGN ITERATIONS TO PERFORMANCE IN ENGINEERING DESIGN","year":2023,"lang":"en","type":"article","venue":"Proceedings of the Design Society","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Engineering design process; Ambiguity; Process (computing); Iterative design; Design process; Industrial engineering; Mathematical optimization; Mathematics; Work in process; Operations management; Engineering","score_opus":0.04541328401669431,"score_gpt":0.2559900809916875,"score_spread":0.2105767969749932,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381185676","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8505247,0.0004687902,0.12592685,0.00081781007,0.00003703811,0.00019344958,0.0000908132,0.0005876337,0.021352923],"genre_scores_gemma":[0.98291636,0.000050548344,0.016429069,0.000019265948,0.000007903766,0.000056434437,0.000041949952,0.000060858725,0.00041776884],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.95006615,0.03145062,0.0032078293,0.003137222,0.010412074,0.0017260859],"domain_scores_gemma":[0.6529384,0.24260077,0.049644224,0.019854307,0.030792287,0.0041699633],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032734133,0.0010522171,0.00046665076,0.0051668365,0.0014579361,0.0060375957,0.0013019212,0.0010435898,0.002250171],"category_scores_gemma":[0.23519838,0.0009422837,0.000510958,0.0022766818,0.0052034757,0.0064220782,0.0042029624,0.002169187,0.0005268162],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012670169,0.00089467375,0.46704063,0.0008786004,0.00034696306,0.00055515644,0.038858946,0.15511341,0.010052163,0.06514466,0.0022011774,0.2576466],"study_design_scores_gemma":[0.00023994353,0.0038821828,0.2731654,0.0010928416,0.00023670199,0.0013793459,0.033874605,0.4513455,0.029461356,0.18412593,0.020621128,0.0005749753],"about_ca_topic_score_codex":0.0012841467,"about_ca_topic_score_gemma":0.0011489951,"teacher_disagreement_score":0.032734133,"about_ca_system_score_codex":0.0029239168,"about_ca_system_score_gemma":0.0016617823,"threshold_uncertainty_score":0.17311674},"labels":[],"label_agreement":null},{"id":"W4381198270","doi":"10.2139/ssrn.4484020","title":"An Empirical Study on Bug Severity Estimation Using Source Code Metrics and Static Analysis","year":2023,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Manitoba; University of Calgary","funders":"","keywords":"Computer science; Leverage (statistics); Software bug; Source lines of code; Java; Code (set theory); Source code; Open source; Domain (mathematical analysis); Data mining; Set (abstract data type); Software; Static analysis; Static program analysis; Code review; Software engineering; Machine learning; Programming language; Software development","score_opus":0.0576716752167631,"score_gpt":0.382528128887947,"score_spread":0.32485645367118393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381198270","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9987288,0.00010834689,0.00067903765,0.000029275916,0.0000038049566,0.000014266526,0.00009465029,0.000012755131,0.0003290191],"genre_scores_gemma":[0.998145,0.00007094815,0.0010658256,0.000015807998,0.0000080192885,0.000017727989,0.00039147882,0.000012658487,0.00027266753],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9947062,0.0028053226,0.0005141515,0.00058233476,0.0011546112,0.00023735384],"domain_scores_gemma":[0.7752476,0.18676431,0.011725741,0.009359741,0.01501542,0.0018873074],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065110703,0.0004825373,0.00035861525,0.0027855306,0.00052887725,0.0007298576,0.00087957527,0.00075643644,0.0016726666],"category_scores_gemma":[0.07448162,0.00033181408,0.0004922795,0.002652893,0.0007825065,0.0021807172,0.0007048287,0.0010521264,0.0005048299],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010168519,0.0043259896,0.93374485,0.00024032025,0.00016368002,0.00032923653,0.0018024397,0.0019017559,0.0033059795,0.0006334102,0.0010427384,0.051492643],"study_design_scores_gemma":[0.00015742588,0.0039760754,0.9548919,0.000112318085,0.0002720241,0.0011440013,0.0029598312,0.030863406,0.0031492368,0.0005260542,0.0019016787,0.000046012985],"about_ca_topic_score_codex":0.0044715838,"about_ca_topic_score_gemma":0.005030826,"teacher_disagreement_score":0.0065110703,"about_ca_system_score_codex":0.00057006895,"about_ca_system_score_gemma":0.0006192349,"threshold_uncertainty_score":0.03443426},"labels":[],"label_agreement":null},{"id":"W4381304075","doi":"10.1109/tse.2023.3281275","title":"Multi-Granularity Detector for Vulnerability Fixes","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Research Foundation Singapore; National University of Singapore","keywords":"Computer science; Commit; Granularity; Vulnerability (computing); Source code; Software; Python (programming language); Code (set theory); Data mining; Computer security; Database; Operating system; Programming language","score_opus":0.03505384984562248,"score_gpt":0.28615123860887126,"score_spread":0.2510973887632488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381304075","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42391136,0.006564897,0.50737244,0.0015563077,0.0010635031,0.0005183207,0.0072880685,0.04458587,0.007139233],"genre_scores_gemma":[0.8544989,0.00074189855,0.13024047,0.00036997432,0.00016422475,0.0001907835,0.008783074,0.00046041858,0.004550248],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9968094,0.00030482645,0.00021214846,0.0011081104,0.0012639454,0.00030152086],"domain_scores_gemma":[0.9939466,0.0021271787,0.0009897514,0.0010015774,0.0015776946,0.00035731943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002675482,0.001494858,0.001298812,0.0082818335,0.00068042916,0.0012452536,0.001928968,0.0017014424,0.0013403945],"category_scores_gemma":[0.010378172,0.000361568,0.0010970378,0.0026930226,0.0005742216,0.0025449235,0.002341325,0.0024117862,0.0012731053],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007344626,0.0006798523,0.12889387,0.0006094997,0.00051970774,0.001105731,0.00037950862,0.07732664,0.028409386,0.0032137127,0.042690378,0.71543735],"study_design_scores_gemma":[0.000034862464,0.0002858284,0.026829813,0.00009263248,0.00015462255,0.0011097324,0.00019232226,0.932365,0.023232246,0.004829755,0.010800124,0.00007305984],"about_ca_topic_score_codex":0.0045066513,"about_ca_topic_score_gemma":0.0070486674,"teacher_disagreement_score":0.0082818335,"about_ca_system_score_codex":0.0009933676,"about_ca_system_score_gemma":0.0012109684,"threshold_uncertainty_score":0.014149487},"labels":[],"label_agreement":null},{"id":"W4381956191","doi":"10.1007/978-3-031-36272-9_6","title":"Investigating the Utility of Self-explanation Through Translation Activities with a Code-Tracing Tutor","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Tracing; Computer science; TUTOR; Code (set theory); TRACE (psycholinguistics); Programming language; Artificial intelligence; Linguistics","score_opus":0.040680946777781155,"score_gpt":0.2753006196332029,"score_spread":0.23461967285542173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381956191","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9776088,0.000041198033,0.015727192,0.00014197422,0.000024130173,0.00032211296,0.000075218944,0.0009381727,0.0051212087],"genre_scores_gemma":[0.98275846,0.000032414107,0.013975241,0.0000557176,0.0000069547273,0.00020785557,0.00014741893,0.00013364949,0.0026823697],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.993804,0.004192724,0.00031281402,0.00074770226,0.00064145646,0.00030144144],"domain_scores_gemma":[0.8793665,0.10440748,0.0032604397,0.006306318,0.0050160964,0.0016431923],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056013106,0.0008101406,0.00060596137,0.00075324485,0.00070734567,0.0025393544,0.0019696981,0.0022519405,0.0063333525],"category_scores_gemma":[0.09396121,0.00046471623,0.0003336832,0.0006147695,0.0007704103,0.0026674212,0.0025294323,0.0012452411,0.0017341346],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012815825,0.020320706,0.11377241,0.0019872251,0.00027981916,0.0032729856,0.13949205,0.019257441,0.14069597,0.0054573254,0.00469963,0.5379486],"study_design_scores_gemma":[0.0052956636,0.04181293,0.13110027,0.0009290625,0.0014723125,0.0031986975,0.058453076,0.435445,0.27858305,0.0146924285,0.028376853,0.00064067385],"about_ca_topic_score_codex":0.0016408844,"about_ca_topic_score_gemma":0.0010626798,"teacher_disagreement_score":0.0063333525,"about_ca_system_score_codex":0.0006324925,"about_ca_system_score_gemma":0.0011368396,"threshold_uncertainty_score":0.029622972},"labels":[],"label_agreement":null},{"id":"W4382363007","doi":"10.1145/3589806.3600043","title":"Fingerprinting and Building Large Reproducible Datasets","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Documentation; Context (archaeology); Data science; Data mining; Software; Scale (ratio); Empirical research; Software quality; Quality (philosophy); Software engineering; Software development; Programming language","score_opus":0.026388828003827226,"score_gpt":0.30869686519255934,"score_spread":0.2823080371887321,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382363007","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051581655,0.0019229103,0.8921241,0.0022123058,0.0007596516,0.001773276,0.033847697,0.011947414,0.0038309589],"genre_scores_gemma":[0.12201856,0.0011105336,0.78058016,0.00074895844,0.0004017206,0.00411457,0.08861921,0.0014256053,0.0009805893],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.93244594,0.02573776,0.010257971,0.015741704,0.014121642,0.0016950131],"domain_scores_gemma":[0.7235565,0.08712204,0.014677156,0.15052511,0.022290273,0.0018289294],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.058681805,0.0018527067,0.0023231802,0.012955055,0.0027341635,0.008695997,0.006188073,0.0034536927,0.0018096407],"category_scores_gemma":[0.22432259,0.0016775419,0.0033599986,0.015647057,0.0026616012,0.01119719,0.010847369,0.003756977,0.0026776078],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011030056,0.0010563452,0.11199974,0.005225108,0.0022581185,0.0013396536,0.003026529,0.09664903,0.035254054,0.0642736,0.07326298,0.60455185],"study_design_scores_gemma":[0.00053592876,0.0008725145,0.05407098,0.002137342,0.00104406,0.0016190424,0.0024683776,0.3293662,0.06638577,0.28242305,0.25855047,0.000526236],"about_ca_topic_score_codex":0.0029580323,"about_ca_topic_score_gemma":0.0039251735,"teacher_disagreement_score":0.9413182,"about_ca_system_score_codex":0.001570804,"about_ca_system_score_gemma":0.0059373896,"threshold_uncertainty_score":0.31034273},"labels":[],"label_agreement":null},{"id":"W4382987313","doi":"10.1145/3607186","title":"What Constitutes the Deployment and Runtime Configuration System? An Empirical Study on OpenStack Projects","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software deployment; Computer science; Configuration Management (ITSM); Software configuration management; Operating system; Software; Leverage (statistics); System deployment; Software system; Software construction","score_opus":0.1711848459207054,"score_gpt":0.38614101168772763,"score_spread":0.21495616576702223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382987313","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9958882,0.00018846852,0.000979735,0.00036007617,0.000008489584,0.000054842916,0.00012876409,0.000022818973,0.0023685803],"genre_scores_gemma":[0.9980568,0.00016182917,0.0009933984,0.00006905436,0.000008880333,0.00006571509,0.00022564772,0.000040577426,0.00037802398],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9816839,0.010393904,0.0013494,0.0021638835,0.0033346447,0.0010741615],"domain_scores_gemma":[0.79061586,0.15101002,0.02859617,0.008070932,0.015699018,0.0060079913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020025196,0.0003549319,0.0003447891,0.0028656977,0.0024279498,0.0047092303,0.0016046873,0.0013960712,0.0023340143],"category_scores_gemma":[0.14775184,0.00053542596,0.00026638084,0.0038496233,0.0045640487,0.012098288,0.0024887493,0.0022261855,0.0004968802],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004575031,0.0010423646,0.72709584,0.00071967894,0.000070746195,0.0011487397,0.18752904,0.0008929113,0.0021925922,0.006861565,0.00482216,0.067166805],"study_design_scores_gemma":[0.000027570164,0.00048404405,0.7494814,0.00046310667,0.000033491713,0.00068191695,0.22462374,0.004087943,0.0009013017,0.0018629928,0.017250443,0.00010211995],"about_ca_topic_score_codex":0.0043965965,"about_ca_topic_score_gemma":0.00557425,"teacher_disagreement_score":0.020025196,"about_ca_system_score_codex":0.0024308157,"about_ca_system_score_gemma":0.0017370981,"threshold_uncertainty_score":0.10590464},"labels":[],"label_agreement":null},{"id":"W4383053810","doi":"10.5281/zenodo.8114639","title":"Investigating the Impact of Pull Request Reviews on Software Quality Metrics and Code Smells","year":2023,"lang":"en","type":"paratext","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code smell; Computer science; Software quality; Quality (philosophy); Code (set theory); Software; Software engineering; Software development; Programming language","score_opus":0.11808314495758414,"score_gpt":0.34749610944964787,"score_spread":0.22941296449206372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383053810","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99307066,0.00030058852,0.0021353078,0.00030907028,0.000057875757,0.00010701228,0.0011774972,0.0007338504,0.0021080088],"genre_scores_gemma":[0.99235266,0.00012513658,0.0031101024,0.000076257886,0.000034054523,0.00009966298,0.0017829447,0.0002716424,0.0021475418],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9870394,0.004583781,0.00070615864,0.0013629313,0.0057570357,0.00055072573],"domain_scores_gemma":[0.6479233,0.26371285,0.03556118,0.013902648,0.034063075,0.004836868],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01088453,0.000714158,0.0005425497,0.0024531996,0.00048707557,0.0018474883,0.0008030231,0.00089490373,0.0034531297],"category_scores_gemma":[0.1597384,0.0003648684,0.0007698657,0.0022500237,0.00049177365,0.0018936971,0.0011514095,0.0011581591,0.0017255349],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0066778376,0.004470377,0.65896183,0.0023493473,0.001301124,0.0012219909,0.0036866497,0.022678565,0.029937036,0.001398886,0.024437899,0.2428785],"study_design_scores_gemma":[0.00025686697,0.007849555,0.88433385,0.0002492705,0.0007164674,0.00044531105,0.0028927533,0.07961082,0.017393038,0.001179486,0.00489877,0.00017376295],"about_ca_topic_score_codex":0.0046761706,"about_ca_topic_score_gemma":0.006361041,"teacher_disagreement_score":0.9891155,"about_ca_system_score_codex":0.0010011061,"about_ca_system_score_gemma":0.0012668468,"threshold_uncertainty_score":0.057563543},"labels":[],"label_agreement":null},{"id":"W4383053811","doi":"10.5281/zenodo.8114660","title":"Investigating the Impact of Pull Request Reviews on Software Quality Metrics and Code Smells","year":2023,"lang":"en","type":"paratext","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code smell; Computer science; Software quality; Quality (philosophy); Code (set theory); Software; Software engineering; Code review; Software development; Programming language","score_opus":0.11808314495758414,"score_gpt":0.34749610944964787,"score_spread":0.22941296449206372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383053811","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99307066,0.00030058852,0.0021353078,0.00030907028,0.000057875757,0.00010701228,0.0011774972,0.0007338504,0.0021080088],"genre_scores_gemma":[0.99235266,0.00012513658,0.0031101024,0.000076257886,0.000034054523,0.00009966298,0.0017829447,0.0002716424,0.0021475418],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9870394,0.004583781,0.00070615864,0.0013629313,0.0057570357,0.00055072573],"domain_scores_gemma":[0.6479233,0.26371285,0.03556118,0.013902648,0.034063075,0.004836868],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01088453,0.000714158,0.0005425497,0.0024531996,0.00048707557,0.0018474883,0.0008030231,0.00089490373,0.0034531297],"category_scores_gemma":[0.1597384,0.0003648684,0.0007698657,0.0022500237,0.00049177365,0.0018936971,0.0011514095,0.0011581591,0.0017255349],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0066778376,0.004470377,0.65896183,0.0023493473,0.001301124,0.0012219909,0.0036866497,0.022678565,0.029937036,0.001398886,0.024437899,0.2428785],"study_design_scores_gemma":[0.00025686697,0.007849555,0.88433385,0.0002492705,0.0007164674,0.00044531105,0.0028927533,0.07961082,0.017393038,0.001179486,0.00489877,0.00017376295],"about_ca_topic_score_codex":0.0046761706,"about_ca_topic_score_gemma":0.006361041,"teacher_disagreement_score":0.9891155,"about_ca_system_score_codex":0.0010011061,"about_ca_system_score_gemma":0.0012668468,"threshold_uncertainty_score":0.057563543},"labels":[],"label_agreement":null},{"id":"W4383067302","doi":"10.1145/3607179","title":"A Systematic Review of Automated Query Reformulations in Source Code Search","year":2023,"lang":"en","type":"review","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada; Dalhousie University","keywords":"Computer science; Information retrieval; Query expansion; Web search query; Vocabulary; Term (time); Weighting; Software; Data mining; Search engine; Programming language","score_opus":0.17950256895553524,"score_gpt":0.4152454057738076,"score_spread":0.23574283681827235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383067302","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020059494,0.9916237,0.0023802675,0.0007307176,0.00013893063,0.0011938413,0.000799968,0.00006979678,0.0010568236],"genre_scores_gemma":[0.01736971,0.9631407,0.013228681,0.0011055791,0.00009411327,0.003123106,0.0015646603,0.000059241243,0.00031418217],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9491001,0.02657936,0.0136892805,0.0023171033,0.007689528,0.0006245845],"domain_scores_gemma":[0.7873412,0.179024,0.012812592,0.004158499,0.015992884,0.0006707929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04711116,0.0023591174,0.006188067,0.03154381,0.0012685824,0.0041101715,0.004326847,0.0021336733,0.00527941],"category_scores_gemma":[0.2089029,0.0017225253,0.006208858,0.030893443,0.0018922772,0.008180022,0.003702886,0.0017781418,0.0011133787],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024019567,0.00004647575,0.00072317285,0.7064533,0.0020920052,0.000119561984,0.0014267216,0.0003319982,0.00045019534,0.0010947805,0.005431553,0.2815901],"study_design_scores_gemma":[0.00031920173,0.00048355578,0.004254841,0.87161624,0.018600902,0.0006003876,0.0016620487,0.000518346,0.0011085256,0.0019290486,0.09878546,0.00012140354],"about_ca_topic_score_codex":0.0102202585,"about_ca_topic_score_gemma":0.02706047,"teacher_disagreement_score":0.04711116,"about_ca_system_score_codex":0.006258567,"about_ca_system_score_gemma":0.026865771,"threshold_uncertainty_score":0.24915063},"labels":[],"label_agreement":null},{"id":"W4383215759","doi":"10.1016/j.jss.2023.111796","title":"A systematic literature review on source code similarity measurement and clone detection: Techniques, applications, and challenges","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Source code; Code review; Field (mathematics); Similarity (geometry); Software; Code (set theory); Java; Software engineering; Static program analysis; Software development; Point (geometry); Data mining; Information retrieval; Programming language; Artificial intelligence","score_opus":0.04170936095502379,"score_gpt":0.2713537760964991,"score_spread":0.2296444151414753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383215759","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002093606,0.9951308,0.00087067805,0.00073014287,0.00015506873,0.0002717415,0.00043839435,0.00001835894,0.00029111878],"genre_scores_gemma":[0.01880752,0.97388625,0.0044647665,0.0012539241,0.00017138374,0.00064100645,0.00063326245,0.00001816715,0.0001238479],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.977091,0.0068852464,0.009138462,0.0020790691,0.004375776,0.00043044015],"domain_scores_gemma":[0.83774316,0.124328196,0.01797962,0.003045959,0.015438856,0.0014641706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028794613,0.0014943143,0.008008717,0.03260463,0.0013718434,0.0043160934,0.0032598784,0.0027276766,0.0024135616],"category_scores_gemma":[0.13002443,0.0013317297,0.0060296943,0.020770833,0.0020224766,0.0058948006,0.0034548235,0.0019448069,0.00044496256],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022354213,0.00005187228,0.003461143,0.8146264,0.00456561,0.00016412644,0.0008315894,0.00016172403,0.0005535056,0.0008271691,0.0036256248,0.17090774],"study_design_scores_gemma":[0.00012424232,0.00036652185,0.00803481,0.9133542,0.033656623,0.0006773888,0.0010628204,0.00020051608,0.0005379321,0.0012616077,0.04064433,0.000078976984],"about_ca_topic_score_codex":0.0077474355,"about_ca_topic_score_gemma":0.030006396,"teacher_disagreement_score":0.03260463,"about_ca_system_score_codex":0.0040592933,"about_ca_system_score_gemma":0.03123831,"threshold_uncertainty_score":0.1522823},"labels":[],"label_agreement":null},{"id":"W4383555717","doi":"10.1145/3607185","title":"Programming by Example Made Easy","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Fundamental Research Funds for the Central Universities; Natural Science Foundation of Jiangsu Province; National Natural Science Foundation of China; Impact Fund; National Science Foundation","keywords":"Computer science; Usability; Domain (mathematical analysis); Table (database); Exploit; Constraint programming; Domain-specific language; Software engineering; Range (aeronautics); Programming language; Theoretical computer science; Data mining; Human–computer interaction; Mathematical optimization","score_opus":0.10711553213041111,"score_gpt":0.3310986989582408,"score_spread":0.22398316682782968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383555717","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006089012,0.00040650088,0.9353393,0.0011641376,0.0002068541,0.00033081777,0.0006030229,0.01313688,0.042723477],"genre_scores_gemma":[0.06440716,0.0007435026,0.90602314,0.000710994,0.00008328006,0.0004505436,0.0016062873,0.003607291,0.022367865],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99702924,0.0009352061,0.00024520047,0.00068566087,0.0009122901,0.00019234775],"domain_scores_gemma":[0.992863,0.0037890451,0.0002982686,0.0020211525,0.00083652535,0.00019195655],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002497311,0.0011586419,0.0006402926,0.00086709816,0.0009375008,0.002741359,0.0022422087,0.0010370901,0.040814526],"category_scores_gemma":[0.014669173,0.00079852826,0.0012877138,0.0008779922,0.0012454707,0.0056747356,0.0035912183,0.0023038161,0.010097653],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028706618,0.00023450203,0.0014155505,0.0012881043,0.00009460193,0.0005352995,0.0014167378,0.0084985215,0.020156376,0.23100646,0.066913076,0.6681537],"study_design_scores_gemma":[0.00010828804,0.00016767406,0.0007105792,0.00042165318,0.00008123614,0.0012824503,0.00040456842,0.055640686,0.02551927,0.18340474,0.73218215,0.00007667532],"about_ca_topic_score_codex":0.0009045368,"about_ca_topic_score_gemma":0.0017794373,"teacher_disagreement_score":0.040814526,"about_ca_system_score_codex":0.0004929002,"about_ca_system_score_gemma":0.0016569517,"threshold_uncertainty_score":0.13653815},"labels":[],"label_agreement":null},{"id":"W4383898381","doi":"10.1109/icse-seip58684.2023.00014","title":"Challenges in Adopting Artificial Intelligence Based User Input Verification Framework in Reporting Software Systems","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Software engineering; Software; Process (computing); Software development","score_opus":0.1769718170743552,"score_gpt":0.3507400884604106,"score_spread":0.1737682713860554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383898381","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016391058,0.0005583079,0.9654065,0.006622948,0.00013866498,0.0006341426,0.000072293864,0.0030964853,0.007079613],"genre_scores_gemma":[0.20211293,0.00038848564,0.7937687,0.0012735275,0.00010049133,0.00034318332,0.00015199895,0.00049594784,0.0013647655],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.8912502,0.059618197,0.008424151,0.006373384,0.030650342,0.0036836679],"domain_scores_gemma":[0.81996965,0.10360655,0.009979999,0.032241635,0.03230707,0.0018950668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.09284499,0.00092320825,0.0011932384,0.002731759,0.0022803347,0.009773141,0.006576244,0.0035020225,0.0021748615],"category_scores_gemma":[0.11985465,0.0016851344,0.002053765,0.0015405987,0.007021148,0.013921625,0.0055822516,0.0076569784,0.0009935971],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025384626,0.0006524707,0.0075794575,0.001609457,0.00020691946,0.0019513322,0.007438254,0.074939765,0.015844455,0.55923414,0.006317043,0.32397294],"study_design_scores_gemma":[0.00017270587,0.0005687919,0.0025547398,0.0019372443,0.000186388,0.0017996742,0.0026757363,0.5396689,0.045054544,0.30171886,0.103335366,0.0003270522],"about_ca_topic_score_codex":0.0074161408,"about_ca_topic_score_gemma":0.006042935,"teacher_disagreement_score":0.09284499,"about_ca_system_score_codex":0.0042953007,"about_ca_system_score_gemma":0.010163683,"threshold_uncertainty_score":0.4910171},"labels":[],"label_agreement":null},{"id":"W4383898397","doi":"10.1109/icse-seip58684.2023.00012","title":"An Empirical Comparison on the Results of Different Clone Detection Setups for C-based Projects","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Huawei Technologies (Canada)","funders":"","keywords":"Code refactoring; clone (Java method); Computer science; Security token; Software maintenance; Source code; Open source; Python (programming language); Code (set theory); Software; Operating system; Programming language; Software system","score_opus":0.10135512877879241,"score_gpt":0.37039760316201775,"score_spread":0.26904247438322537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383898397","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9617389,0.0041474667,0.02070802,0.00022979344,0.00014085922,0.00036918264,0.0020927272,0.007717289,0.0028556683],"genre_scores_gemma":[0.9411449,0.0010627281,0.04868671,0.00011452075,0.00006337424,0.0003059463,0.006433026,0.0009611117,0.0012276551],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.96657306,0.010139479,0.004551334,0.0061381836,0.011220673,0.0013772579],"domain_scores_gemma":[0.72076565,0.18425587,0.023308232,0.025825975,0.041609798,0.0042344755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016118929,0.001970312,0.0012309186,0.007584703,0.0011565099,0.0022631385,0.0016692413,0.0017656358,0.0007581714],"category_scores_gemma":[0.13321409,0.0006704914,0.0010997154,0.005145569,0.0013387926,0.004438735,0.0022663677,0.0017834218,0.0008071435],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0060937763,0.0032131663,0.2714144,0.007187479,0.0018078809,0.0007904674,0.0055124317,0.030335527,0.09259762,0.0016949255,0.015996443,0.5633558],"study_design_scores_gemma":[0.00065875263,0.015486832,0.54827654,0.001172653,0.0019621982,0.0041848533,0.005554641,0.2073414,0.18756144,0.0022765133,0.024694452,0.00082974124],"about_ca_topic_score_codex":0.0026436525,"about_ca_topic_score_gemma":0.0035989934,"teacher_disagreement_score":0.016118929,"about_ca_system_score_codex":0.0012606292,"about_ca_system_score_gemma":0.0011097441,"threshold_uncertainty_score":0.085246086},"labels":[],"label_agreement":null},{"id":"W4384009618","doi":"10.1109/msr59073.2023.00033","title":"On Codex Prompt Engineering for OCL Generation: An Empirical Study","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Science and Engineering Research Council; Canadian Institute for Advanced Research","keywords":"Computer science; Programming language; Object Constraint Language; Task (project management); Unified Modeling Language; Natural language processing; Syntax; Artificial intelligence; Object (grammar); Software engineering; UML tool; Software","score_opus":0.09103536348884478,"score_gpt":0.3692618199990915,"score_spread":0.27822645651024674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384009618","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9515589,0.005868117,0.019168954,0.0011477086,0.00027070186,0.0010272816,0.005445648,0.007567267,0.007945368],"genre_scores_gemma":[0.9314341,0.0015406235,0.03750714,0.00092351495,0.00013306989,0.0009965823,0.02157214,0.002455831,0.0034371598],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9651813,0.0188398,0.002517151,0.0050772154,0.007764514,0.00062008423],"domain_scores_gemma":[0.49984616,0.4357339,0.017120982,0.020781059,0.02453126,0.0019865748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026320713,0.0016056753,0.00087558123,0.0040950202,0.0015952187,0.0035455902,0.0031073315,0.0027416016,0.002680283],"category_scores_gemma":[0.34416085,0.00088426465,0.0010946962,0.0034969454,0.002467544,0.007084528,0.0034227709,0.003820493,0.0020150268],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044512805,0.0033730217,0.28064895,0.010932223,0.0008535417,0.0030183783,0.031383853,0.029699313,0.011814032,0.0045519685,0.06456532,0.5547081],"study_design_scores_gemma":[0.0012527573,0.0058060307,0.41277942,0.0055639083,0.0013206162,0.0069858893,0.020610576,0.33137244,0.027607203,0.009681881,0.17618668,0.0008326814],"about_ca_topic_score_codex":0.01144845,"about_ca_topic_score_gemma":0.008830916,"teacher_disagreement_score":0.026320713,"about_ca_system_score_codex":0.0021270239,"about_ca_system_score_gemma":0.0021717004,"threshold_uncertainty_score":0.1391989},"labels":[],"label_agreement":null},{"id":"W4384009674","doi":"10.1109/msr59073.2023.00057","title":"Evolution of the Practice of Software Testing in Java Projects","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Java; Computer science; Software engineering; Software quality; Software; Test (biology); Regression testing; Software bug; Software testing; Source lines of code; Open source software; Software development; Programming language; Software construction","score_opus":0.036807040282827996,"score_gpt":0.29524027052692325,"score_spread":0.25843323024409526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384009674","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98442185,0.0012973223,0.0061424375,0.0024333254,0.00004663028,0.000070627655,0.00010252026,0.00005284407,0.005432419],"genre_scores_gemma":[0.9948978,0.00048210024,0.0038497446,0.00018620062,0.000022624574,0.000043519893,0.00010013838,0.0000246594,0.00039318655],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9563645,0.019707916,0.0027344124,0.005401022,0.013972112,0.0018200743],"domain_scores_gemma":[0.7183478,0.12988056,0.06680204,0.019082554,0.05812435,0.0077627837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029002724,0.00034042526,0.00033403997,0.0049491273,0.0013317886,0.0041678054,0.0014568709,0.001215734,0.0008056529],"category_scores_gemma":[0.16401911,0.00050143996,0.0004192922,0.0049772486,0.0026520845,0.0048824716,0.003080882,0.0019873995,0.00028349983],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014387075,0.0003241743,0.8097063,0.00026325756,0.000101410464,0.00050475,0.019953286,0.0016138564,0.001873243,0.0038279612,0.0014701991,0.16021775],"study_design_scores_gemma":[0.00002628125,0.00048015008,0.9563795,0.00048094476,0.00006490952,0.001163058,0.0150004905,0.0054329005,0.0012736386,0.0032560695,0.016354444,0.0000877638],"about_ca_topic_score_codex":0.007969765,"about_ca_topic_score_gemma":0.009747554,"teacher_disagreement_score":0.029002724,"about_ca_system_score_codex":0.0045927777,"about_ca_system_score_gemma":0.0039271615,"threshold_uncertainty_score":0.1533829},"labels":[],"label_agreement":null},{"id":"W4384009705","doi":"10.1109/msr59073.2023.00075","title":"PyMigBench: A Benchmark for Python Library Migration","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Python (programming language); Computer science; Java; Popularity; World Wide Web; Programming language; Software engineering","score_opus":0.023764875170130127,"score_gpt":0.2675182896619241,"score_spread":0.24375341449179397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384009705","genre_codex":"software","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37329158,0.0059014214,0.09651632,0.0019875981,0.0017336989,0.0014118127,0.07642987,0.3892748,0.053452894],"genre_scores_gemma":[0.433985,0.002549308,0.2189085,0.0009597179,0.00017627866,0.0017741902,0.26067135,0.067284115,0.013691574],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9922529,0.0013249974,0.0010445414,0.0010100188,0.0034276722,0.00093980035],"domain_scores_gemma":[0.9885689,0.0034157995,0.000919243,0.0031924387,0.0030494565,0.0008542696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004842651,0.002635716,0.0008859851,0.0037658252,0.0015396388,0.0016696167,0.005074928,0.0011174722,0.0042902096],"category_scores_gemma":[0.02197083,0.0012438484,0.0013486274,0.0075138127,0.0013332373,0.0042545083,0.0033451112,0.0028563333,0.0039354186],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032472075,0.0015989596,0.046567284,0.006356047,0.0008606656,0.0015333653,0.001937583,0.080135904,0.03230998,0.015169042,0.5376056,0.27267846],"study_design_scores_gemma":[0.0010006516,0.0013658452,0.06802088,0.000806823,0.0003990477,0.0023165755,0.0013337743,0.34209812,0.12043206,0.019777551,0.44187942,0.00056925503],"about_ca_topic_score_codex":0.011790697,"about_ca_topic_score_gemma":0.01129778,"teacher_disagreement_score":0.011790697,"about_ca_system_score_codex":0.0017942957,"about_ca_system_score_gemma":0.003823352,"threshold_uncertainty_score":0.025610685},"labels":[],"label_agreement":null},{"id":"W4384009718","doi":"10.1109/icse-nier58687.2023.00007","title":"CodeS: Towards Code Model Generalization Under Distribution Shift","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"JST-Mirai Program; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Abstract syntax tree; Python (programming language); Source code; Programming language; Java; Code (set theory); Code refactoring; Source lines of code; Syntax; Artificial intelligence; Theoretical computer science; Software","score_opus":0.04480406771770869,"score_gpt":0.3074395312124491,"score_spread":0.26263546349474043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384009718","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20236763,0.0015989539,0.75649095,0.0029284384,0.0003484118,0.00026980057,0.004705998,0.026601959,0.004687895],"genre_scores_gemma":[0.77083844,0.0006954519,0.19419329,0.002067819,0.0002371907,0.00048696922,0.022984259,0.0025121223,0.0059844255],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99678445,0.0007336573,0.00015753086,0.0013729889,0.00066698215,0.00028434518],"domain_scores_gemma":[0.98950785,0.0041630715,0.00068475626,0.0036958084,0.0015757605,0.00037270586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004256371,0.0021165502,0.0012102679,0.0020504703,0.0008446577,0.0022503748,0.003142788,0.0023992942,0.0019557471],"category_scores_gemma":[0.026688986,0.0007419618,0.0018223753,0.0014397989,0.0016852338,0.004856823,0.0041906866,0.0057087974,0.0021008793],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008144894,0.0004300842,0.021743715,0.00045264993,0.00032424112,0.00032526313,0.00046976187,0.43086562,0.013227681,0.018869031,0.051106434,0.46137106],"study_design_scores_gemma":[0.000033313292,0.000078989804,0.00089736754,0.000025819154,0.000022489176,0.00007730275,0.000058932685,0.96621734,0.004228431,0.025169756,0.003174269,0.000015967817],"about_ca_topic_score_codex":0.010384256,"about_ca_topic_score_gemma":0.01010303,"teacher_disagreement_score":0.010384256,"about_ca_system_score_codex":0.0024389788,"about_ca_system_score_gemma":0.0030580608,"threshold_uncertainty_score":0.022510111},"labels":[],"label_agreement":null},{"id":"W4384009747","doi":"10.1109/icse-companion58688.2023.00069","title":"A Framework to Communicate Software Engineering Data Effectively with Dashboards","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Dashboard; Data science; Data visualization; Visualization; Presentation (obstetrics); Process (computing); Software; Set (abstract data type); Face (sociological concept); Software engineering; World Wide Web; Data mining","score_opus":0.04970416740297849,"score_gpt":0.3156825463589012,"score_spread":0.2659783789559227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384009747","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002085839,0.00039621958,0.97890526,0.0028559766,0.00017324711,0.00029625368,0.00021709944,0.0019923171,0.0130778095],"genre_scores_gemma":[0.06713476,0.0009984893,0.9229601,0.00052538066,0.00016693275,0.00083741447,0.00059315533,0.00036467757,0.0064190514],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9858576,0.00798542,0.0016577207,0.0017207232,0.0020450372,0.00073355966],"domain_scores_gemma":[0.982169,0.008967071,0.0016410823,0.00307056,0.0027789394,0.0013733853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01992837,0.0023520722,0.0010665456,0.010048441,0.004059083,0.017316753,0.0051464764,0.0055386648,0.008007826],"category_scores_gemma":[0.023971627,0.0016495482,0.0028111113,0.0068414183,0.012404709,0.02550085,0.01083828,0.006232316,0.00319934],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030083469,0.00007831917,0.0005629835,0.00024931805,0.000032501903,0.0004007013,0.0062483544,0.0034195837,0.0011129408,0.95402974,0.0034026967,0.03043275],"study_design_scores_gemma":[0.000055224406,0.00015879961,0.00058672123,0.0013114921,0.00009454681,0.0005465686,0.005107009,0.053343087,0.0026744166,0.6222955,0.31366345,0.0001631488],"about_ca_topic_score_codex":0.009427825,"about_ca_topic_score_gemma":0.0069732196,"teacher_disagreement_score":0.01992837,"about_ca_system_score_codex":0.0029944386,"about_ca_system_score_gemma":0.0058344537,"threshold_uncertainty_score":0.105392516},"labels":[],"label_agreement":null},{"id":"W4384009757","doi":"10.1109/msr59073.2023.00054","title":"An Empirical Study to Investigate Collaboration Among Developers in Open Source Software (OSS)","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Documentation; Python (programming language); Software engineering; Source code; World Wide Web; Code review; Empirical research; Software; Open source software; Acknowledgement; Public domain software; Software development; Static program analysis; Programming language; Computer security","score_opus":0.04972716717945835,"score_gpt":0.3668609772587295,"score_spread":0.31713381007927116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384009757","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99371356,0.00016873916,0.0021979485,0.00025947703,0.000012516458,0.0001944474,0.0004959619,0.00001388411,0.0029433784],"genre_scores_gemma":[0.9937785,0.00018328827,0.0038113822,0.00012920365,0.000024924837,0.0005213545,0.0008251336,0.000013916278,0.0007123099],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98593616,0.008725288,0.0010689661,0.0011112925,0.002294102,0.0008642445],"domain_scores_gemma":[0.8622772,0.08975867,0.022028552,0.0052481866,0.0136853345,0.007002031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011778337,0.0003079374,0.00034334027,0.003564764,0.0022395006,0.0016929697,0.0007756848,0.0008653339,0.001917021],"category_scores_gemma":[0.0636898,0.0002999962,0.00032680167,0.0042725387,0.0010921494,0.0030959449,0.0024693392,0.0011986843,0.0006480335],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023213506,0.0017658129,0.9082257,0.00059013785,0.00008381213,0.00042447317,0.053137384,0.0004502243,0.00092933135,0.002124084,0.002192854,0.029843932],"study_design_scores_gemma":[0.00010083293,0.0012868834,0.8163,0.00069316034,0.0000900311,0.00072088925,0.1514444,0.00508251,0.0018093451,0.0021361555,0.02026729,0.000068512825],"about_ca_topic_score_codex":0.002682576,"about_ca_topic_score_gemma":0.004664919,"teacher_disagreement_score":0.011778337,"about_ca_system_score_codex":0.0013586256,"about_ca_system_score_gemma":0.0025958032,"threshold_uncertainty_score":0.06229055},"labels":[],"label_agreement":null},{"id":"W4384026501","doi":"10.1109/msr59073.2023.00085","title":"Defectors: A Large, Diverse Python Dataset for Defect Prediction","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Python (programming language); Computer science; Approx; Machine learning; Artificial intelligence; Source lines of code; Source code; Data mining; Programming language; Software; Operating system","score_opus":0.06861649982258129,"score_gpt":0.32963252042095736,"score_spread":0.2610160205983761,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026501","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11354402,0.000565284,0.015205213,0.0008602719,0.00022325804,0.00044168858,0.82449746,0.03896041,0.0057023913],"genre_scores_gemma":[0.075448915,0.00024071061,0.01640607,0.00022005459,0.00004768114,0.0004646986,0.9035211,0.001399436,0.0022513089],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99818546,0.00019321356,0.00020167857,0.0003630675,0.00082820246,0.00022841716],"domain_scores_gemma":[0.9955201,0.0008248826,0.0006058628,0.0013720852,0.001166119,0.0005109996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012824141,0.0016062275,0.00065074424,0.0036145996,0.0007465841,0.00078268885,0.0022689244,0.0013400338,0.003686301],"category_scores_gemma":[0.0067929877,0.000495394,0.0010491662,0.003854283,0.0008436788,0.0017383372,0.0020514042,0.001954438,0.0050634155],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010053185,0.0009185289,0.06635253,0.0014691151,0.00023010663,0.0012266993,0.00041732838,0.01729124,0.00938968,0.0031514298,0.8134291,0.08511894],"study_design_scores_gemma":[0.00082370004,0.0009708845,0.2392643,0.00045389845,0.00021974332,0.0030165347,0.0007663391,0.16326591,0.03279707,0.0155482255,0.54244065,0.00043278633],"about_ca_topic_score_codex":0.010247166,"about_ca_topic_score_gemma":0.016711464,"teacher_disagreement_score":0.010247166,"about_ca_system_score_codex":0.0009160871,"about_ca_system_score_gemma":0.0022296114,"threshold_uncertainty_score":0.020375073},"labels":[],"label_agreement":null},{"id":"W4384026505","doi":"10.1109/msr59073.2023.00067","title":"DACOS—A Manually Annotated Dataset of Code Smells","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Code smell; Code (set theory); Context (archaeology); Information retrieval; Benchmarking; Artificial intelligence; Focus (optics); Machine learning; Natural language processing; World Wide Web; Software","score_opus":0.03287142682747221,"score_gpt":0.31310794928898916,"score_spread":0.280236522461517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026505","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10384801,0.0018750419,0.015764473,0.00057332683,0.00037145184,0.00044060455,0.84508926,0.02312479,0.008913101],"genre_scores_gemma":[0.04619895,0.00041655116,0.026765926,0.00018848162,0.000052367814,0.00056171615,0.92105645,0.0013670019,0.0033925949],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99641526,0.000483726,0.00053344807,0.00091859914,0.0014191638,0.00022972022],"domain_scores_gemma":[0.9831759,0.0046521877,0.0023477229,0.00297329,0.006068885,0.0007820598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016005901,0.0015646345,0.0007187729,0.008043788,0.0011105156,0.0010980619,0.0014055128,0.001819432,0.003751473],"category_scores_gemma":[0.013042042,0.00046522415,0.0008188316,0.006199817,0.00072964694,0.0019343531,0.0015730873,0.0015164288,0.0050757374],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017090703,0.0006660691,0.05079519,0.006160488,0.0002448626,0.0016152553,0.0019518986,0.006930387,0.037653405,0.0034604894,0.7138473,0.1749656],"study_design_scores_gemma":[0.00031269438,0.0005713145,0.13637449,0.00097246165,0.00015351328,0.0018604547,0.0014013815,0.0355271,0.03911233,0.0046737203,0.7786646,0.0003760396],"about_ca_topic_score_codex":0.010600244,"about_ca_topic_score_gemma":0.028872345,"teacher_disagreement_score":0.010600244,"about_ca_system_score_codex":0.0011891706,"about_ca_system_score_gemma":0.0018925834,"threshold_uncertainty_score":0.021077096},"labels":[],"label_agreement":null},{"id":"W4384026519","doi":"10.1109/msr59073.2023.00060","title":"Do Subjectivity and Objectivity Always Agreeƒ A Case Study with Stack Overflow Questions","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Objectivity (philosophy); Subjectivity; Voting; Quality (philosophy); Reliability (semiconductor); Artificial intelligence; Machine learning; Mechanism (biology); Epistemology","score_opus":0.02586928310656097,"score_gpt":0.29452650959381294,"score_spread":0.268657226487252,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026519","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9339377,0.0009494654,0.037786093,0.0035143357,0.00010772912,0.00062857056,0.00030464114,0.00018535432,0.022586074],"genre_scores_gemma":[0.9948297,0.00008611428,0.003908245,0.00030522683,0.00005849779,0.000115942115,0.000098472105,0.000059650476,0.00053816114],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.8713503,0.09194261,0.0069158277,0.009503114,0.017691052,0.0025970344],"domain_scores_gemma":[0.28402057,0.6198626,0.044471964,0.019461468,0.029426267,0.0027571209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.12585692,0.00066879764,0.0010353082,0.0034438302,0.0022759247,0.005895798,0.0013339784,0.0026541208,0.002301623],"category_scores_gemma":[0.4326859,0.0006234504,0.00087331067,0.001957409,0.005537892,0.011159878,0.0037945944,0.0020958958,0.00049287366],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015363182,0.0004885285,0.59797835,0.0017808495,0.00073154777,0.0013963146,0.22047082,0.0032106705,0.007640101,0.020348217,0.0041196393,0.1402987],"study_design_scores_gemma":[0.00038470665,0.0017092366,0.6520319,0.0019742688,0.00086156727,0.0019942594,0.14113365,0.04915962,0.013871256,0.08547443,0.05081895,0.0005861921],"about_ca_topic_score_codex":0.003108828,"about_ca_topic_score_gemma":0.00250594,"teacher_disagreement_score":0.12585692,"about_ca_system_score_codex":0.0028465658,"about_ca_system_score_gemma":0.0018913317,"threshold_uncertainty_score":0.665603},"labels":[],"label_agreement":null},{"id":"W4384026542","doi":"10.1109/icse-companion58688.2023.00077","title":"Towards Utilizing Natural Language Processing Techniques to Assist in Software Engineering Tasks","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Executable; Source code; Code review; Natural language; Natural language processing; Artificial intelligence; Code (set theory); Codebase; Programming language; Task (project management); Embedding; Software; Software development; Static program analysis; Engineering","score_opus":0.01798308103686853,"score_gpt":0.3021009793877318,"score_spread":0.28411789835086326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026542","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005788056,0.00034623937,0.98790514,0.0011049727,0.000045294244,0.0001559453,0.00028894158,0.0033578551,0.0010076562],"genre_scores_gemma":[0.036219973,0.00043458532,0.9607073,0.0003180304,0.00005359076,0.00019719913,0.0010260029,0.00023553686,0.000807713],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9940268,0.0029798439,0.000392176,0.0013756175,0.0010922609,0.00013326407],"domain_scores_gemma":[0.97537917,0.016339004,0.0020568145,0.0033751077,0.0025596444,0.00029034584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054551084,0.0024977885,0.0008525874,0.004199569,0.0006081391,0.0028768745,0.002154861,0.0016869723,0.0028120633],"category_scores_gemma":[0.029248059,0.00079874095,0.0016047896,0.00266843,0.0019664564,0.009527749,0.0030058331,0.004242332,0.0026339716],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017970074,0.00063226797,0.0040535135,0.0014571128,0.0001429102,0.00031737128,0.0015229385,0.047925774,0.030514294,0.035475817,0.012744562,0.8650338],"study_design_scores_gemma":[0.000066931374,0.00020375643,0.0015566206,0.00031828505,0.000094489704,0.00034691123,0.00085017295,0.7364427,0.023400743,0.2075373,0.029100971,0.00008108087],"about_ca_topic_score_codex":0.002028528,"about_ca_topic_score_gemma":0.0037570589,"teacher_disagreement_score":0.0054551084,"about_ca_system_score_codex":0.0009653176,"about_ca_system_score_gemma":0.0025552623,"threshold_uncertainty_score":0.028849721},"labels":[],"label_agreement":null},{"id":"W4384026586","doi":"10.1109/msr59073.2023.00023","title":"Evaluating Software Documentation Quality","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Documentation; Readability; Computer science; World Wide Web; Software documentation; JavaScript; Python (programming language); Software; Quality (philosophy); Java; Software engineering; Software development; Software development process; Programming language","score_opus":0.1352538161411699,"score_gpt":0.456957610908366,"score_spread":0.3217037947671961,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026586","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92997545,0.001778008,0.053686596,0.0006477468,0.0001373412,0.0009819143,0.0010390917,0.0016036482,0.010150327],"genre_scores_gemma":[0.908239,0.0008455998,0.08534432,0.00014721083,0.000060098104,0.000671525,0.0021687637,0.0003285277,0.0021949783],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9451041,0.01637797,0.008604826,0.00243586,0.026136588,0.0013405881],"domain_scores_gemma":[0.66314894,0.16370517,0.044387907,0.0118014775,0.11198485,0.004971639],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04402148,0.0008713163,0.00095875293,0.010874149,0.00092569354,0.002962571,0.0009625619,0.001020794,0.0017684764],"category_scores_gemma":[0.2054072,0.0004585736,0.0010178428,0.0056557264,0.0007937435,0.004009008,0.0022247655,0.0009055458,0.0007027783],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082672184,0.0011658374,0.3518098,0.0030162332,0.0006011536,0.0002776492,0.012014014,0.005279807,0.016578782,0.0022310284,0.008177631,0.5980213],"study_design_scores_gemma":[0.0002442628,0.0070133945,0.8354435,0.0025048987,0.00078586786,0.001128825,0.013332941,0.047677938,0.04853901,0.004205218,0.03853824,0.0005858817],"about_ca_topic_score_codex":0.0022904333,"about_ca_topic_score_gemma":0.0035120223,"teacher_disagreement_score":0.9559785,"about_ca_system_score_codex":0.0020678057,"about_ca_system_score_gemma":0.0023271844,"threshold_uncertainty_score":0.23281056},"labels":[],"label_agreement":null},{"id":"W4384026638","doi":"10.1109/icse-nier58687.2023.00031","title":"How does quality deviate in stable releases by backporting?","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Context (archaeology); Computer science; Upgrade; Quality (philosophy); Process (computing); Software bug; Software evolution; Code (set theory); Stability (learning theory); Software quality; Software maintenance; Software; Outlier; Software engineering; Risk analysis (engineering); Software development; Business; Programming language; Software construction; Artificial intelligence; Operating system","score_opus":0.031118806501731863,"score_gpt":0.3071146140783392,"score_spread":0.2759958075766073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026638","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9503084,0.00087758276,0.025134072,0.0058441414,0.00015159864,0.00009855102,0.0005151575,0.00082285353,0.016247611],"genre_scores_gemma":[0.9960121,0.00015272669,0.0022398145,0.00029051866,0.000032969732,0.000016661801,0.00015834258,0.00014710138,0.0009497647],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98915803,0.0021319515,0.00073161954,0.0023782866,0.004578024,0.0010220108],"domain_scores_gemma":[0.915881,0.029457359,0.02725871,0.012033001,0.012151967,0.003217998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008289454,0.0004436579,0.00043236327,0.0025741372,0.0012553254,0.0058468473,0.0012377105,0.0016703518,0.0033759035],"category_scores_gemma":[0.1006339,0.0005294023,0.00047782605,0.0023819741,0.0031975855,0.008639974,0.0026466413,0.002969804,0.0009256124],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007134243,0.00040185382,0.73265153,0.00047945115,0.00024654155,0.0030850542,0.025326751,0.0068393275,0.013568797,0.026636736,0.0070145433,0.18303601],"study_design_scores_gemma":[0.000050742958,0.00093183084,0.8798423,0.00044051447,0.00017924898,0.0023103473,0.022881882,0.021383852,0.0080652,0.04137937,0.022343123,0.00019161975],"about_ca_topic_score_codex":0.0051005366,"about_ca_topic_score_gemma":0.0044081793,"teacher_disagreement_score":0.008289454,"about_ca_system_score_codex":0.0022543685,"about_ca_system_score_gemma":0.0011027277,"threshold_uncertainty_score":0.043839335},"labels":[],"label_agreement":null},{"id":"W4384026673","doi":"10.1109/msr59073.2023.00058","title":"Keep the Ball Rolling: Analyzing Release Cadence in GitHub Projects","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Carleton University","funders":"","keywords":"Cadence; Computer science; Java; Leverage (statistics); Software; Python (programming language); Software engineering; Programming language; World Wide Web; Artificial intelligence; Engineering","score_opus":0.037754048155799164,"score_gpt":0.28612667247671014,"score_spread":0.24837262432091098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026673","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98520976,0.0012394345,0.005268968,0.00019325077,0.000031850832,0.00010128093,0.0049416535,0.0005322865,0.0024815204],"genre_scores_gemma":[0.97577685,0.00060685485,0.00897385,0.00006957791,0.00006547345,0.0003441183,0.012535273,0.00034227266,0.0012857132],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99478936,0.0014075419,0.0007183099,0.0010654144,0.0015963493,0.00042308404],"domain_scores_gemma":[0.93865776,0.03587131,0.01384031,0.0032659594,0.006646504,0.0017181371],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061298944,0.00049115624,0.0004588191,0.010229273,0.00067005225,0.001961796,0.00064905675,0.000478254,0.00093002897],"category_scores_gemma":[0.047019128,0.00033733965,0.0005475231,0.008386308,0.00059856294,0.0024049515,0.0017716046,0.000773043,0.0006143197],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041272913,0.00014014026,0.8653164,0.0010064669,0.00020125229,0.0005080711,0.011818839,0.0015009376,0.0038179597,0.0009026915,0.007500388,0.10687418],"study_design_scores_gemma":[0.000015028638,0.00010877524,0.9773072,0.00017579603,0.00007959071,0.0003945459,0.0050462782,0.00542739,0.0012164308,0.0007290605,0.009446534,0.00005343579],"about_ca_topic_score_codex":0.0056047877,"about_ca_topic_score_gemma":0.0072603435,"teacher_disagreement_score":0.010229273,"about_ca_system_score_codex":0.0006285995,"about_ca_system_score_gemma":0.0007330013,"threshold_uncertainty_score":0.03241837},"labels":[],"label_agreement":null},{"id":"W4384026759","doi":"10.1109/msr59073.2023.00015","title":"Understanding the Time to First Response in GitHub Pull Requests","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Codebase; Computer science; Response time; Process (computing); Productivity; Exploratory research; Empirical research; Software; Operating system","score_opus":0.0825497501348978,"score_gpt":0.2959970350030007,"score_spread":0.21344728486810288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384026759","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9840947,0.00038931714,0.011126957,0.00037938196,0.000032520773,0.00012827522,0.00045762918,0.00030949095,0.003081875],"genre_scores_gemma":[0.99460334,0.00012507377,0.0035029168,0.00006283217,0.000021676498,0.00008459607,0.00046912083,0.000097316646,0.0010330318],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99302095,0.0024719823,0.0007052745,0.0012158834,0.0017665196,0.0008192992],"domain_scores_gemma":[0.8872714,0.0731128,0.02422999,0.003945933,0.008516367,0.0029235044],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007616386,0.0005908723,0.0004892565,0.0032098514,0.0007935593,0.002551957,0.00096857006,0.0012232955,0.0026797138],"category_scores_gemma":[0.06880548,0.0005447758,0.00044935764,0.0016608783,0.00082177453,0.004174155,0.0015580804,0.0014114241,0.0014652504],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010972117,0.0005373296,0.85976106,0.00068514777,0.00017635497,0.0016865389,0.014468112,0.0067299916,0.019277053,0.004649941,0.0029447456,0.08798646],"study_design_scores_gemma":[0.00004422316,0.00057593203,0.91769755,0.00017526867,0.00013309409,0.0018941342,0.014043809,0.04365723,0.0069600544,0.004327116,0.01032163,0.00017009264],"about_ca_topic_score_codex":0.0047363085,"about_ca_topic_score_gemma":0.004570613,"teacher_disagreement_score":0.007616386,"about_ca_system_score_codex":0.001033589,"about_ca_system_score_gemma":0.0013774419,"threshold_uncertainty_score":0.040279806},"labels":[],"label_agreement":null},{"id":"W4384039146","doi":"10.1007/s10664-023-10342-7","title":"BTLink : automatic link recovery between issues and commits based on pre-trained BERT model","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Science Foundation of Jiangsu Province; National Natural Science Foundation of China","keywords":"Commit; Computer science; Traceability; Identifier; Software; Software engineering; Classifier (UML); Data mining; Machine learning; Artificial intelligence; Database; Programming language","score_opus":0.035326283152714626,"score_gpt":0.3095404327340525,"score_spread":0.2742141495813379,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384039146","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10902684,0.0015504307,0.6652384,0.0012634018,0.0012209797,0.00046096518,0.007782446,0.20786718,0.00558937],"genre_scores_gemma":[0.6647153,0.0004520453,0.277741,0.0007173288,0.0003929984,0.00042155656,0.028407821,0.0043125106,0.02283936],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986092,0.00019505218,0.00007984149,0.00047104014,0.00046276406,0.00018206576],"domain_scores_gemma":[0.9948651,0.0020262145,0.0004149918,0.0013185748,0.0010122501,0.00036288096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017881214,0.0019540398,0.0011167949,0.00275898,0.0007584828,0.0015498401,0.0034171427,0.0022606829,0.009170246],"category_scores_gemma":[0.010211979,0.00077130797,0.0011090507,0.0014065673,0.0004560624,0.004315712,0.0021623848,0.0034783366,0.007935707],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021460983,0.0014413932,0.018981997,0.00053729466,0.00032995737,0.00064178376,0.00021431765,0.12913181,0.022082854,0.0039258837,0.13038598,0.6901808],"study_design_scores_gemma":[0.00004257039,0.00009551092,0.0011692649,0.000017534834,0.000029356976,0.00005938297,0.000029963981,0.987586,0.0043272646,0.0034476724,0.0031764824,0.000019137575],"about_ca_topic_score_codex":0.012933578,"about_ca_topic_score_gemma":0.020960782,"teacher_disagreement_score":0.012933578,"about_ca_system_score_codex":0.0008777447,"about_ca_system_score_gemma":0.0024043096,"threshold_uncertainty_score":0.030677497},"labels":[],"label_agreement":null},{"id":"W4384158949","doi":"10.1109/icpc58990.2023.00033","title":"UnityLint: A Bad Smell Detector for Unity","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code smell; Computer science; Metadata; Leverage (statistics); Video game; Software; Source code; Video game development; Domain (mathematical analysis); Human–computer interaction; Graphics; Software engineering; Open source; World Wide Web; Multimedia; Software development; Game design; Software quality; Artificial intelligence; Computer graphics (images); Programming language","score_opus":0.04147912389474392,"score_gpt":0.30086318579584326,"score_spread":0.25938406190109936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384158949","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14821644,0.0019523519,0.3603368,0.0011333622,0.00040635534,0.00088103226,0.032814357,0.43669286,0.017566355],"genre_scores_gemma":[0.4671126,0.0009787756,0.42318147,0.0010928073,0.000112414746,0.0008573949,0.06234542,0.0240302,0.020288896],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9973763,0.00026104093,0.0003189656,0.00062541745,0.0012505408,0.00016762078],"domain_scores_gemma":[0.9904108,0.0039758035,0.0022379672,0.0013622531,0.0017112711,0.0003019238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016980569,0.0016650447,0.00066912686,0.006920361,0.0005214394,0.0020915272,0.0013179898,0.0014352403,0.0046510817],"category_scores_gemma":[0.015348397,0.0007298305,0.0009059833,0.0017864181,0.0006804901,0.003165255,0.002454062,0.0010561474,0.004591529],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013308097,0.00037721414,0.100184284,0.0028804187,0.0002849255,0.0033398992,0.003091081,0.0044364636,0.061177183,0.008529349,0.20284483,0.6115237],"study_design_scores_gemma":[0.00024313004,0.00085761683,0.12618606,0.0013078966,0.00035671028,0.0077086394,0.0018739898,0.23097523,0.23104618,0.02209744,0.3767276,0.00061955187],"about_ca_topic_score_codex":0.0029860213,"about_ca_topic_score_gemma":0.0047798203,"teacher_disagreement_score":0.006920361,"about_ca_system_score_codex":0.00074099266,"about_ca_system_score_gemma":0.0009589859,"threshold_uncertainty_score":0.015559375},"labels":[],"label_agreement":null},{"id":"W4384159047","doi":"10.1109/icpc58990.2023.00013","title":"PyVerDetector: A Chrome Extension Detecting the Python Version of Stack Overflow Code Snippets","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Japan Society for the Promotion of Science","keywords":"Python (programming language); Computer science; Programming language; Source code","score_opus":0.02462400889865477,"score_gpt":0.27712107428689625,"score_spread":0.25249706538824146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384159047","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10477948,0.0012540398,0.18916748,0.0010548821,0.000867441,0.000780771,0.021510275,0.6695179,0.011067691],"genre_scores_gemma":[0.4293512,0.0012538772,0.36707643,0.002748649,0.00042776513,0.0012976582,0.03420811,0.1289739,0.034662396],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99727505,0.00028676208,0.0002531989,0.00069246447,0.0012318371,0.00026066444],"domain_scores_gemma":[0.9843439,0.008633236,0.0018732965,0.0017423648,0.0029171607,0.0004901259],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002790291,0.0021293021,0.00093945366,0.003604478,0.0008132937,0.0020099508,0.0014019723,0.0016586162,0.016187457],"category_scores_gemma":[0.023495272,0.0011228534,0.0009986317,0.0013667529,0.0011774191,0.0040510856,0.0027737683,0.001444641,0.0072911824],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004194199,0.00058012566,0.066008255,0.0032481875,0.00036643626,0.007690791,0.003812908,0.0040916735,0.082650684,0.009073337,0.36861712,0.44966626],"study_design_scores_gemma":[0.00048716864,0.0007675016,0.09082542,0.001482513,0.0002668158,0.00783806,0.0012294959,0.07087314,0.35694894,0.012942487,0.45523262,0.0011057648],"about_ca_topic_score_codex":0.0025939138,"about_ca_topic_score_gemma":0.0037142178,"teacher_disagreement_score":0.016187457,"about_ca_system_score_codex":0.00061176607,"about_ca_system_score_gemma":0.0013293112,"threshold_uncertainty_score":0.05415243},"labels":[],"label_agreement":null},{"id":"W4384159091","doi":"10.1109/icpc58990.2023.00031","title":"Pathways to Leverage Transcompiler based Data Augmentation for Cross-Language Clone Detection","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; clone (Java method); Leverage (statistics); Source code; Exploit; Programming language; Software; Artificial intelligence","score_opus":0.09413039926532048,"score_gpt":0.3641596172484101,"score_spread":0.27002921798308965,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384159091","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1040908,0.001577662,0.84816194,0.0015940722,0.0005249716,0.00025288083,0.002801952,0.036456026,0.0045397314],"genre_scores_gemma":[0.42292002,0.0006292138,0.5543224,0.0014013337,0.00011846865,0.00050174334,0.010632302,0.002512393,0.006961999],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984491,0.00032548726,0.00010137409,0.0006734994,0.00033809218,0.00011239338],"domain_scores_gemma":[0.99500567,0.0014248508,0.00039086991,0.001620141,0.0013796713,0.0001788209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018110237,0.0018416625,0.0008212743,0.001481824,0.00060846953,0.0014187739,0.0023798246,0.0015471464,0.0034911449],"category_scores_gemma":[0.01122989,0.0006596992,0.0013596993,0.0012519241,0.00097378914,0.004098872,0.0029786485,0.003114989,0.0027746563],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010165472,0.00065060786,0.022770727,0.000640951,0.0002767579,0.0010522738,0.0010137479,0.079532616,0.094140075,0.007613376,0.01940947,0.7718829],"study_design_scores_gemma":[0.00006973615,0.00031770972,0.003033444,0.00009968125,0.00012509711,0.00063523074,0.00023162333,0.846821,0.10489236,0.021066824,0.02261933,0.0000880017],"about_ca_topic_score_codex":0.0035588108,"about_ca_topic_score_gemma":0.0058434834,"teacher_disagreement_score":0.0035588108,"about_ca_system_score_codex":0.0007718274,"about_ca_system_score_gemma":0.00155072,"threshold_uncertainty_score":0.011678994},"labels":[],"label_agreement":null},{"id":"W4384162682","doi":"10.1145/3605158.3605849","title":"Towards Reliable Memory Management for Python Native Extensions","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); University of New Brunswick","funders":"","keywords":"Python (programming language); Computer science; Programming language; Garbage collection; Operating system; Byte; Garbage","score_opus":0.03569893466496222,"score_gpt":0.3087346104936955,"score_spread":0.2730356758287333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384162682","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08004795,0.0015824913,0.7615463,0.0015509169,0.00072146585,0.00042699295,0.00061453425,0.13534492,0.018164448],"genre_scores_gemma":[0.4519775,0.00093260035,0.5030957,0.0021625205,0.00037971156,0.00096014864,0.0024680267,0.019360868,0.018662963],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99471575,0.0008830497,0.00047963858,0.0005820877,0.002413669,0.0009257478],"domain_scores_gemma":[0.98814046,0.001476862,0.0008769591,0.005833454,0.0030389456,0.0006333105],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055379737,0.0014427623,0.00084476435,0.0016173759,0.001513953,0.0039023666,0.006709829,0.0013157285,0.004442996],"category_scores_gemma":[0.018996043,0.0014445453,0.000964345,0.0013048683,0.0019551886,0.008384131,0.008424417,0.0038869828,0.0029536604],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002402722,0.0006635153,0.022318702,0.0019507984,0.00029875935,0.0014852835,0.0034948732,0.028006429,0.10160206,0.11895517,0.15079686,0.5680249],"study_design_scores_gemma":[0.0003162976,0.00058954157,0.004954933,0.00062931556,0.0002884029,0.001249709,0.00080019527,0.3720248,0.2647006,0.08989088,0.26413774,0.00041752204],"about_ca_topic_score_codex":0.003016069,"about_ca_topic_score_gemma":0.0035371096,"teacher_disagreement_score":0.006709829,"about_ca_system_score_codex":0.0015020474,"about_ca_system_score_gemma":0.0042620637,"threshold_uncertainty_score":0.029287934},"labels":[],"label_agreement":null},{"id":"W4384302745","doi":"10.1109/icse48619.2023.00214","title":"CoLeFunDa: Explainable Silent Vulnerability Fix Identification","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Huawei Technologies (Canada)","funders":"National Key Research and Development Program of China","keywords":"Leverage (statistics); Computer science; Vulnerability (computing); Exploit; Identification (biology); Computer security; Function (biology); Vulnerability assessment; Commit; Artificial intelligence; Psychology","score_opus":0.03163195279413376,"score_gpt":0.3007841687061718,"score_spread":0.26915221591203803,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384302745","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.072666414,0.0016758019,0.8428436,0.001268223,0.00014503002,0.0004091532,0.005068782,0.07261162,0.0033114527],"genre_scores_gemma":[0.47552323,0.000479462,0.50395364,0.00080633484,0.000092633796,0.00059067085,0.010302266,0.0015543444,0.0066973795],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835205,0.00040832144,0.00007564887,0.00064609625,0.00037627062,0.00014159754],"domain_scores_gemma":[0.9941288,0.0035116319,0.00050239573,0.0011715414,0.00053328794,0.00015238152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022657844,0.0020226452,0.000932736,0.0022300314,0.000555133,0.0011849677,0.0034571392,0.002798005,0.0051540285],"category_scores_gemma":[0.014079598,0.0006541681,0.0015130788,0.00068198633,0.0013651138,0.0035794503,0.0031793448,0.0035094449,0.0017237057],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007530811,0.000557903,0.031047734,0.00088387897,0.0002704636,0.0011819834,0.0013242081,0.10872876,0.021452473,0.014690703,0.040358502,0.77875036],"study_design_scores_gemma":[0.00005694855,0.00019799409,0.0038840594,0.000108687746,0.000059722162,0.00048161158,0.00016022862,0.9412424,0.01583336,0.020620806,0.017278709,0.000075605116],"about_ca_topic_score_codex":0.004975707,"about_ca_topic_score_gemma":0.0132488115,"teacher_disagreement_score":0.0051540285,"about_ca_system_score_codex":0.0011710966,"about_ca_system_score_gemma":0.001864415,"threshold_uncertainty_score":0.017242014},"labels":[],"label_agreement":null},{"id":"W4384302749","doi":"10.1109/icse48619.2023.00205","title":"Retrieval-Based Prompt Selection for Code-Related Few-Shot Learning","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":156,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Assertion; Task (project management); Code (set theory); Embedding; Artificial intelligence; Selection (genetic algorithm); Language model; Source code; Natural language processing; Programming language; Machine learning","score_opus":0.038158346489101116,"score_gpt":0.30657942551462763,"score_spread":0.26842107902552653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384302749","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09085193,0.0031803825,0.8488328,0.00063178554,0.0004164106,0.0006209887,0.0014169369,0.05058655,0.0034622448],"genre_scores_gemma":[0.62583023,0.00077306497,0.3497419,0.0016768579,0.0002741676,0.0011047653,0.009659218,0.0015559058,0.009383859],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982535,0.00048826748,0.00009350217,0.00080205366,0.00021524263,0.00014732532],"domain_scores_gemma":[0.993087,0.004431169,0.00024976942,0.0010025642,0.000891901,0.00033752728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029792986,0.0025562746,0.0018876098,0.0016357858,0.0007517697,0.0012497373,0.0038525425,0.002734305,0.004747191],"category_scores_gemma":[0.013425039,0.0008762831,0.0013552094,0.0011184409,0.0010588781,0.0043679527,0.0024064118,0.0038310674,0.003921268],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011662665,0.0012755928,0.0047579478,0.0010313991,0.00022711889,0.00044841925,0.00054651854,0.13130382,0.026502607,0.0036653103,0.031707555,0.7973674],"study_design_scores_gemma":[0.00013802016,0.0004304018,0.0008547493,0.000042746862,0.00006472206,0.00016771992,0.00013511651,0.9733314,0.011706884,0.009543784,0.0035336206,0.000050742823],"about_ca_topic_score_codex":0.0048891115,"about_ca_topic_score_gemma":0.008598667,"teacher_disagreement_score":0.0048891115,"about_ca_system_score_codex":0.0011844976,"about_ca_system_score_gemma":0.0020886976,"threshold_uncertainty_score":0.015880883},"labels":[],"label_agreement":null},{"id":"W4384302756","doi":"10.1109/icse48619.2023.00159","title":"Semi-Automatic, Inline and Collaborative Web Page Code Curations","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung","keywords":"Computer science; World Wide Web; Source code; Web page; Code (set theory); Software; Code review; Web application; Web development; Web browser; Information retrieval; Software quality; Software development; Software engineering; The Internet; Programming language; Set (abstract data type)","score_opus":0.02036975070914896,"score_gpt":0.29717672495841424,"score_spread":0.27680697424926526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384302756","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25224075,0.000322802,0.686989,0.00024650127,0.00008730773,0.001170829,0.0011305698,0.052600104,0.0052120485],"genre_scores_gemma":[0.3368167,0.0001992769,0.65084565,0.00014134204,0.000050912255,0.00074327964,0.0021189197,0.0022090466,0.0068749147],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9933357,0.0028422885,0.00047792963,0.001487257,0.0016469675,0.0002098576],"domain_scores_gemma":[0.9283871,0.03706357,0.0073656645,0.017366359,0.00821248,0.001604844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004605813,0.0013792503,0.000987322,0.0026209971,0.0010235138,0.0022844458,0.0023436898,0.0012378854,0.0029070405],"category_scores_gemma":[0.03508304,0.0010103986,0.0005367187,0.0012410365,0.00076689577,0.0028642372,0.0036063744,0.0010941416,0.003056577],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000997389,0.0020817367,0.022486938,0.0016182027,0.00024172221,0.0007301822,0.016944878,0.0074539087,0.15759729,0.0015715497,0.016095024,0.77218115],"study_design_scores_gemma":[0.00083237374,0.0038052902,0.11738018,0.0006098628,0.0005189272,0.0033224947,0.00884301,0.33237916,0.37565452,0.013263247,0.1424478,0.0009431513],"about_ca_topic_score_codex":0.0023603847,"about_ca_topic_score_gemma":0.0066756406,"teacher_disagreement_score":0.004605813,"about_ca_system_score_codex":0.00044630794,"about_ca_system_score_gemma":0.0018543447,"threshold_uncertainty_score":0.024358153},"labels":[],"label_agreement":null},{"id":"W4384302785","doi":"10.1109/icse48619.2023.00063","title":"Explaining Software Bugs Leveraging Code Structures in Neural Machine Translation","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Dalhousie University","keywords":"Computer science; Software bug; Leverage (statistics); Source code; Machine translation; Code review; Software; Code (set theory); Artificial intelligence; Software engineering; Security bug; Machine learning; Programming language; Static program analysis; Software development; Software security assurance; Computer security","score_opus":0.08106133440923918,"score_gpt":0.3153356945573667,"score_spread":0.23427436014812753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384302785","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16562814,0.0025427754,0.7952611,0.002170851,0.00032291017,0.0003333759,0.002166387,0.025020625,0.006553796],"genre_scores_gemma":[0.72663665,0.00096902234,0.26129562,0.0005296465,0.0001263838,0.00019917313,0.005536374,0.0009115093,0.0037956566],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990361,0.0003676876,0.000057361816,0.00028861122,0.00019091215,0.00005929397],"domain_scores_gemma":[0.9934076,0.0050703874,0.00034871532,0.0005882969,0.000502724,0.00008226123],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012514365,0.001210947,0.0004970263,0.0024139073,0.00056407147,0.0009538734,0.0012920155,0.0018064437,0.0037994734],"category_scores_gemma":[0.011536222,0.0005699155,0.0013036546,0.0016019378,0.00071292685,0.0022379765,0.0011977147,0.0017594729,0.001766993],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003931376,0.0003656535,0.010835708,0.0009488167,0.0001712181,0.0015651849,0.0011868139,0.1897986,0.015835399,0.015912663,0.020649046,0.74233776],"study_design_scores_gemma":[0.00005396845,0.00009067265,0.0011469232,0.00006686872,0.00006145607,0.00033260562,0.0001120351,0.96054167,0.006329042,0.02694675,0.0042899367,0.000028226916],"about_ca_topic_score_codex":0.006133088,"about_ca_topic_score_gemma":0.012119353,"teacher_disagreement_score":0.006133088,"about_ca_system_score_codex":0.001088347,"about_ca_system_score_gemma":0.001333638,"threshold_uncertainty_score":0.012710512},"labels":[],"label_agreement":null},{"id":"W4384302787","doi":"10.1109/icse48619.2023.00108","title":"Code Review of Build System Specifications: Prevalence, Purposes, Patterns, and Perceptions","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code review; Computer science; Executable; Software engineering; Context (archaeology); Code refactoring; Software quality; Static program analysis; Empirical research; Software system; Code (set theory); Software development; Software; Programming language; Set (abstract data type)","score_opus":0.04411486717807045,"score_gpt":0.29943261221828193,"score_spread":0.2553177450402115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384302787","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9814966,0.0024749516,0.0073012845,0.0034864317,0.000097065385,0.00011121514,0.00007700254,0.00018590532,0.0047695087],"genre_scores_gemma":[0.99352187,0.0011917702,0.0029817848,0.00070805696,0.000052348823,0.000076823206,0.00013702587,0.00013755166,0.0011927747],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9336622,0.029234236,0.0046972195,0.0032777216,0.027299808,0.0018288001],"domain_scores_gemma":[0.51374197,0.3309892,0.078761905,0.012629974,0.05599873,0.00787815],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035994805,0.00036725504,0.0005156779,0.006942878,0.0024584562,0.0040422603,0.0015163436,0.0017591848,0.0009555539],"category_scores_gemma":[0.2925462,0.0007926516,0.00042997205,0.003087723,0.003948264,0.0046317754,0.004289688,0.0022767629,0.00032659236],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002704084,0.00014883836,0.34106883,0.0014645407,0.000121244295,0.0022385085,0.45459425,0.00050255854,0.011939342,0.0022821575,0.005451313,0.17991796],"study_design_scores_gemma":[0.000055160803,0.0008384078,0.49616688,0.002539112,0.00014690998,0.00795238,0.38198385,0.005954351,0.007865153,0.0029686713,0.09309994,0.0004291589],"about_ca_topic_score_codex":0.008439563,"about_ca_topic_score_gemma":0.011074463,"teacher_disagreement_score":0.035994805,"about_ca_system_score_codex":0.0040922905,"about_ca_system_score_gemma":0.005405277,"threshold_uncertainty_score":0.19036102},"labels":[],"label_agreement":null},{"id":"W4384345670","doi":"10.1109/icse48619.2023.00157","title":"Demystifying Issues, Challenges, and Solutions for Multilingual Software Development","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Interoperability; Computer science; Interfacing; Key (lock); Variety (cybernetics); World Wide Web; Software engineering; Software development; Software; Data science; Programming language; Computer security; Artificial intelligence","score_opus":0.12685703829467435,"score_gpt":0.33720996416821486,"score_spread":0.2103529258735405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384345670","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8031049,0.008604626,0.06475236,0.103971265,0.0006526935,0.00041851596,0.00020898825,0.0006937472,0.017592862],"genre_scores_gemma":[0.950576,0.002284406,0.041141883,0.0025121877,0.00016090724,0.00019167396,0.00014845356,0.00022285391,0.0027616732],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9789184,0.01266914,0.0014892119,0.0016778857,0.003462508,0.0017829398],"domain_scores_gemma":[0.93495816,0.034997556,0.010914528,0.0032470555,0.011368659,0.004514012],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02735929,0.00089164317,0.0005354872,0.0066142646,0.008808682,0.011803377,0.0016424545,0.0024492596,0.0022066957],"category_scores_gemma":[0.06094402,0.00076594413,0.00076879683,0.0045186165,0.007394508,0.029361997,0.009014246,0.003697324,0.00047765666],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010394764,0.00022020255,0.0832069,0.0013777677,0.000056562316,0.0036920323,0.6546975,0.0007583578,0.007041812,0.04351834,0.009061487,0.19626515],"study_design_scores_gemma":[0.00002005344,0.00012333076,0.032516167,0.0013871721,0.000057758003,0.001813546,0.80615395,0.0038424085,0.0025902963,0.045554668,0.10575598,0.00018470263],"about_ca_topic_score_codex":0.007855384,"about_ca_topic_score_gemma":0.017987706,"teacher_disagreement_score":0.02735929,"about_ca_system_score_codex":0.0061723576,"about_ca_system_score_gemma":0.013903406,"threshold_uncertainty_score":0.14469147},"labels":[],"label_agreement":null},{"id":"W4384913059","doi":"10.1007/978-981-19-9948-2_9","title":"Intelligent Software Maintenance","year":2023,"lang":"en","type":"book-chapter","venue":"Natural computing series","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; Polytechnique Montréal","funders":"","keywords":"Computer science; Software development; Software maintenance; Software engineering; Software construction; Software system; Software; Leverage (statistics); Context (archaeology); Artificial intelligence; Software analytics; Programming language","score_opus":0.023219877670793136,"score_gpt":0.2597016250617113,"score_spread":0.23648174739091815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384913059","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041484055,0.017346073,0.23516758,0.0026307884,0.0011018814,0.00008783636,0.00022092169,0.0027836727,0.7365129],"genre_scores_gemma":[0.04705443,0.009497279,0.09043439,0.00079846656,0.00060991937,0.00008149292,0.00069104985,0.00060901965,0.850224],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9998338,0.000019639809,0.000006245397,0.000037859827,0.000092392096,0.000010039493],"domain_scores_gemma":[0.9998086,0.00006920855,0.000010895039,0.000053102463,0.000047250174,0.00001111815],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00021799105,0.0005783599,0.00035951583,0.000866198,0.0004444026,0.0013092932,0.00070623966,0.0006030518,0.030834347],"category_scores_gemma":[0.0006546052,0.0002892579,0.00027581767,0.0008396203,0.0006208293,0.001886006,0.0005875976,0.0009282221,0.014860948],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017495608,0.00004177937,0.00014848054,0.00016065447,0.000009265883,0.00003544444,0.0001566817,0.0020967156,0.0021905967,0.093117364,0.16235937,0.7396661],"study_design_scores_gemma":[0.0000072498246,0.000038087503,0.00068299635,0.00018016613,0.000017290742,0.00030179688,0.00007133633,0.014660455,0.0031041882,0.13701396,0.8439082,0.000014341913],"about_ca_topic_score_codex":0.0006844501,"about_ca_topic_score_gemma":0.0016674711,"teacher_disagreement_score":0.030834347,"about_ca_system_score_codex":0.00053155504,"about_ca_system_score_gemma":0.00040089778,"threshold_uncertainty_score":0.10315114},"labels":[],"label_agreement":null},{"id":"W4385015322","doi":"10.48550/arxiv.2307.10236","title":"Look Before You Leap: An Exploratory Study of Uncertainty Measurement for Large Language Models","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"JST-Mirai Program; Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Trustworthiness; Estimation; Computer science; Computer security; Economics","score_opus":0.17216525414989875,"score_gpt":0.244593140594098,"score_spread":0.07242788644419926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385015322","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85031056,0.0013196638,0.14057067,0.002275387,0.00005736082,0.00043353997,0.0011942902,0.0011821495,0.0026563876],"genre_scores_gemma":[0.9274985,0.0001715715,0.06997803,0.0002904627,0.000037653506,0.00025000077,0.0011865788,0.00024251666,0.00034460545],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96936494,0.023799729,0.0011142131,0.0020423683,0.0032706878,0.00040796972],"domain_scores_gemma":[0.50864893,0.46254212,0.008553146,0.013741986,0.0055382065,0.0009755762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031461027,0.0008810477,0.00065809,0.0021126943,0.0011741829,0.002506672,0.0019642622,0.0015868942,0.000938327],"category_scores_gemma":[0.21996799,0.0005745706,0.0010589352,0.0018064056,0.0024323345,0.005738533,0.003189286,0.004074125,0.0002694992],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027091494,0.0021073506,0.2498857,0.002681288,0.0012288352,0.0028668127,0.046295047,0.2892392,0.016085789,0.058855437,0.015544228,0.31250116],"study_design_scores_gemma":[0.00013367503,0.0008075609,0.026326442,0.00027120317,0.00013036653,0.00085912005,0.0049043563,0.9008011,0.009294544,0.04730482,0.008980416,0.00018626885],"about_ca_topic_score_codex":0.004405742,"about_ca_topic_score_gemma":0.0053081075,"teacher_disagreement_score":0.031461027,"about_ca_system_score_codex":0.0018736637,"about_ca_system_score_gemma":0.0013195424,"threshold_uncertainty_score":0.1663838},"labels":[],"label_agreement":null},{"id":"W4385075908","doi":"10.1016/j.jss.2023.111806","title":"Study the correlation between the readme file of GitHub projects and their popularity","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of Manitoba","funders":"","keywords":"Popularity; Computer science; License; World Wide Web; Database; Point (geometry); Information retrieval; Operating system; Mathematics","score_opus":0.049228368613494264,"score_gpt":0.2802505586975987,"score_spread":0.2310221900841044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385075908","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9957028,0.00018306667,0.00058712205,0.00014280339,0.000015794018,0.000013990854,0.0018044467,0.000126311,0.0014236845],"genre_scores_gemma":[0.9966176,0.00006585103,0.0004187517,0.000021336986,0.00002479287,0.00001550692,0.0014661424,0.000055105505,0.0013149424],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99757785,0.00054979936,0.00028797842,0.00050275406,0.0008371739,0.000244385],"domain_scores_gemma":[0.8555861,0.09489134,0.029648906,0.004726741,0.011248321,0.003898572],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016136956,0.00024109874,0.00029654006,0.00289589,0.0003410449,0.0013189454,0.0005291775,0.0006838125,0.0067550666],"category_scores_gemma":[0.058095563,0.00022673619,0.00030711293,0.0033438618,0.00043902837,0.002067598,0.0007527382,0.0010841881,0.0023386902],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026571302,0.000090512396,0.9886953,0.000060700953,0.00007326339,0.00009629493,0.0003859865,0.0002279465,0.00060664554,0.000107644315,0.0008952827,0.008494704],"study_design_scores_gemma":[0.000007986747,0.00018063627,0.9947318,0.000013378066,0.0000346776,0.00021399415,0.00072957005,0.0023699682,0.00068963366,0.00012163658,0.00089045434,0.000016420765],"about_ca_topic_score_codex":0.0029452217,"about_ca_topic_score_gemma":0.0048045786,"teacher_disagreement_score":0.0067550666,"about_ca_system_score_codex":0.00035555402,"about_ca_system_score_gemma":0.00045429278,"threshold_uncertainty_score":0.022597909},"labels":[],"label_agreement":null},{"id":"W4385241228","doi":"10.1080/07366981.2023.2229986","title":"A maturity level assessment of the use of technology by internal audit functions: a comparative analysis of the Federal Government of Canada","year":2023,"lang":"en","type":"article","venue":"EDPACS","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Maturity (psychological); Internal audit; Audit; Government (linguistics); Business; Work (physics); Capability Maturity Model; Empirical research; Accounting; Public relations; Political science; Engineering; Computer science; Software","score_opus":0.04985382864943331,"score_gpt":0.29500594321496115,"score_spread":0.24515211456552785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385241228","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9755706,0.001169394,0.0006196205,0.0011329703,0.000011783453,0.00010818711,0.0011392405,0.000037293525,0.020210955],"genre_scores_gemma":[0.99637043,0.00059559965,0.0006165444,0.000063301086,0.0000019247839,0.000017382683,0.00052943575,0.000010522698,0.0017949439],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9912525,0.0005834793,0.0002906276,0.00043880614,0.004999149,0.002435452],"domain_scores_gemma":[0.96263564,0.0025818092,0.002527278,0.000686038,0.028331224,0.0032381227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046815383,0.00026201963,0.00032728777,0.00964184,0.0051977593,0.0050840275,0.0011213927,0.00046979933,0.0012672782],"category_scores_gemma":[0.016765168,0.00028851567,0.0004674432,0.015469421,0.0017246412,0.0016380878,0.0019740635,0.0009127259,0.00013958523],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":true,"study_design_scores_codex":[0.00018212864,0.00010683872,0.84216934,0.00020088007,0.00008105846,0.00040720485,0.039139155,0.0014410235,0.0016360796,0.013051692,0.004148353,0.09743619],"study_design_scores_gemma":[0.0000035009318,0.00004553533,0.9447386,0.00013007883,0.00002641313,0.000072892624,0.040535916,0.001115069,0.0005301628,0.00019090375,0.012565348,0.000045663393],"about_ca_topic_score_codex":0.99440515,"about_ca_topic_score_gemma":0.9958734,"teacher_disagreement_score":0.99440515,"about_ca_system_score_codex":0.14634073,"about_ca_system_score_gemma":0.14698625,"threshold_uncertainty_score":0.99012375},"labels":[],"label_agreement":null},{"id":"W4385270039","doi":"10.1109/nlbse59153.2023.00019","title":"Evaluating Code Comment Generation With Summarized API Docs","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Documentation; Application programming interface; Java; Leverage (statistics); Code (set theory); Artificial intelligence; Machine learning; Programming language","score_opus":0.1338888887134159,"score_gpt":0.37542359165664413,"score_spread":0.24153470294322824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385270039","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7154999,0.0053643677,0.16031712,0.0016021973,0.0008643395,0.0013662247,0.021049522,0.08553032,0.008405913],"genre_scores_gemma":[0.7376888,0.00088027585,0.18105109,0.00046992704,0.0002041777,0.00051371794,0.07271022,0.0011427664,0.005339079],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99717665,0.0010408881,0.00024375376,0.00067089725,0.0007338109,0.00013390098],"domain_scores_gemma":[0.98274314,0.012241059,0.0007578321,0.001562788,0.0023101997,0.0003849512],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040386626,0.002338512,0.0009678888,0.0029940968,0.0004820307,0.0014966355,0.0016655998,0.0019252488,0.0023649333],"category_scores_gemma":[0.022249622,0.00033398203,0.0011730783,0.0015212384,0.00044559542,0.0023582848,0.00086519006,0.0011616069,0.0021963469],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023852717,0.0017203733,0.026434569,0.002217504,0.0007036964,0.0007055127,0.00044198812,0.34421253,0.016023645,0.0013507447,0.060673073,0.5431311],"study_design_scores_gemma":[0.00014859246,0.0005752743,0.0035232215,0.0000620369,0.00010513977,0.00013718984,0.00014226469,0.9762793,0.013602407,0.00074179267,0.0046420954,0.000040605795],"about_ca_topic_score_codex":0.011594335,"about_ca_topic_score_gemma":0.0144767575,"teacher_disagreement_score":0.011594335,"about_ca_system_score_codex":0.0014642342,"about_ca_system_score_gemma":0.0015086605,"threshold_uncertainty_score":0.023053706},"labels":[],"label_agreement":null},{"id":"W4385325593","doi":"10.1109/gas59301.2023.00009","title":"An Exploratory Approach for Game Engine Architecture Recovery","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Concordia University","funders":"","keywords":"Computer science; Plug-in; Rendering (computer graphics); Game programming; Game Developer; Game development tool; Game design; Architecture; Video game development; Graphics; Software engineering; Game art design; Game design document; Human–computer interaction; Operating system; Artificial intelligence","score_opus":0.029375938140579233,"score_gpt":0.2681597852795064,"score_spread":0.23878384713892717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385325593","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03444758,0.00012263368,0.94986916,0.0008558789,0.000046542405,0.0007485987,0.00038027475,0.008338914,0.005190532],"genre_scores_gemma":[0.1199404,0.000088833105,0.8725229,0.00023858505,0.000014206668,0.00063020556,0.0010854289,0.0011873768,0.004292046],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9941164,0.0021086878,0.0002848431,0.0010576404,0.0019520352,0.00048033116],"domain_scores_gemma":[0.9851327,0.0061212913,0.0010891032,0.0051208893,0.0022464606,0.00028951862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058272947,0.0022020107,0.0008885204,0.0039632865,0.001955894,0.0031548324,0.0036570828,0.0022132285,0.0047170483],"category_scores_gemma":[0.023306629,0.0016677212,0.0025219554,0.0013942815,0.0020931521,0.004185389,0.004975728,0.0036441789,0.0018863282],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007992815,0.0018934327,0.029895851,0.0017641226,0.000401858,0.0028529787,0.025030946,0.06990853,0.11174782,0.14570695,0.025574941,0.58442324],"study_design_scores_gemma":[0.00012609696,0.0006825324,0.006457525,0.00053621904,0.0002184773,0.0018274669,0.0074625737,0.7230093,0.06409753,0.091688685,0.1036653,0.00022838588],"about_ca_topic_score_codex":0.0038012192,"about_ca_topic_score_gemma":0.009523048,"teacher_disagreement_score":0.0058272947,"about_ca_system_score_codex":0.0014474356,"about_ca_system_score_gemma":0.0034138884,"threshold_uncertainty_score":0.030818045},"labels":[],"label_agreement":null},{"id":"W4385452280","doi":"10.1109/sds57534.2023.00022","title":"Automated Extraction of IoT Critical Objects from IoT Storylines, Requirements and User Stories via NLP","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Internet of Things; Process (computing); Artificial intelligence; Task (project management); Machine learning; Natural language processing; World Wide Web; Programming language; Systems engineering","score_opus":0.03868667499743109,"score_gpt":0.3490463881627202,"score_spread":0.3103597131652891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385452280","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22133577,0.0019738555,0.63664186,0.0028965056,0.00042224894,0.0022157063,0.07494116,0.033632908,0.025939943],"genre_scores_gemma":[0.36792555,0.0010715496,0.515944,0.00037825154,0.00011127418,0.0012340996,0.10547634,0.0009521207,0.0069067893],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984682,0.00063879014,0.00015190765,0.00037372147,0.000294014,0.00007331847],"domain_scores_gemma":[0.991699,0.005423766,0.0007853573,0.0005502454,0.0014075586,0.00013413794],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001512265,0.0019334033,0.00043364367,0.0027575437,0.0005816086,0.0011179759,0.0008557066,0.0010840854,0.0050319103],"category_scores_gemma":[0.008684763,0.00050091615,0.00093281595,0.0012912325,0.00054055697,0.0034660264,0.001279037,0.0013563922,0.0032936085],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071516825,0.00047301938,0.015692936,0.0066502877,0.00023499622,0.007142359,0.009353876,0.032109223,0.073346585,0.013392563,0.121983856,0.71890515],"study_design_scores_gemma":[0.000098444645,0.0003548703,0.035405517,0.00113622,0.0003379789,0.0031949482,0.012622576,0.5250462,0.109064214,0.018924253,0.29356226,0.00025249727],"about_ca_topic_score_codex":0.0050096,"about_ca_topic_score_gemma":0.007761028,"teacher_disagreement_score":0.0050319103,"about_ca_system_score_codex":0.0010783004,"about_ca_system_score_gemma":0.001342521,"threshold_uncertainty_score":0.016833425},"labels":[],"label_agreement":null},{"id":"W4385478004","doi":"10.1109/compsac57700.2023.00091","title":"MaGnn: Binary-Source Code Matching by Modality-Sharing Graph Convolution for Binary Provenance Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Binary number; Provenance; Mathematics; Arithmetic; Geology; Petrology","score_opus":0.02261132786985806,"score_gpt":0.28533141829153674,"score_spread":0.26272009042167865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385478004","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028680567,0.00021935221,0.949133,0.00018388842,0.000065051354,0.00011768678,0.0006719223,0.019139003,0.0017895359],"genre_scores_gemma":[0.43382806,0.00020896629,0.5549263,0.00034489666,0.000058651247,0.00023386748,0.004041637,0.0013662423,0.0049914173],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9994136,0.00006717529,0.000031403462,0.00020339714,0.00022780463,0.00005656592],"domain_scores_gemma":[0.99932194,0.00014619913,0.000111121706,0.00020133462,0.00017818811,0.000041172203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005736571,0.00070610974,0.0005281569,0.0020770787,0.0004415666,0.00071042427,0.0014133683,0.00083853415,0.0031634937],"category_scores_gemma":[0.0034301062,0.00029818324,0.0007554301,0.0010136357,0.00065886386,0.001765304,0.0016331403,0.00088774494,0.0010439701],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046374483,0.00021819065,0.005544398,0.0002775961,0.00010914606,0.0004659179,0.00024550603,0.07915059,0.051789407,0.01919886,0.014422444,0.8281142],"study_design_scores_gemma":[0.000016533579,0.000037392194,0.0012138104,0.000013426023,0.000015569112,0.0001498269,0.00003156712,0.96097684,0.015906716,0.018305892,0.003313315,0.000019067818],"about_ca_topic_score_codex":0.008397684,"about_ca_topic_score_gemma":0.011279164,"teacher_disagreement_score":0.008397684,"about_ca_system_score_codex":0.0011173513,"about_ca_system_score_gemma":0.001217862,"threshold_uncertainty_score":0.016697645},"labels":[],"label_agreement":null},{"id":"W4385489212","doi":"10.1109/compsac57700.2023.00112","title":"Prediction of Bug Inducing Commits Using Metrics Trend Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; IBM (Canada); Western University","funders":"","keywords":"Commit; Computer science; Metric (unit); Process (computing); Quality (philosophy); Source code; Software quality; Software development; Software; Programming language; Database; Engineering","score_opus":0.10481882265082362,"score_gpt":0.31376973137565795,"score_spread":0.20895090872483432,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385489212","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94166714,0.0009856502,0.038866017,0.00033410155,0.00010213337,0.00012877888,0.012098965,0.004278739,0.0015384596],"genre_scores_gemma":[0.9404684,0.00035649928,0.033346668,0.00003279726,0.000066920955,0.00013244577,0.024210185,0.00025074513,0.0011354652],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9985494,0.00018739625,0.00016038331,0.00044579332,0.0005442927,0.00011284634],"domain_scores_gemma":[0.97645116,0.009336947,0.0064142076,0.0020153492,0.0047772033,0.0010050425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028719844,0.00095748954,0.0005329245,0.009773025,0.00034347255,0.00092883833,0.00065274123,0.0007102536,0.0006427569],"category_scores_gemma":[0.018730402,0.0003192331,0.00059133995,0.00413085,0.00022711785,0.0013137759,0.00068193855,0.0010803888,0.00090364687],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030506082,0.00020255202,0.84344995,0.00024029597,0.00013776222,0.0003444427,0.00031695332,0.017871778,0.005590524,0.0009125827,0.009107816,0.12152031],"study_design_scores_gemma":[0.000048475846,0.00063003914,0.52718395,0.00007900905,0.000104784456,0.000715001,0.00044145004,0.45262024,0.0076724994,0.0028169907,0.0076163677,0.00007124795],"about_ca_topic_score_codex":0.0052526887,"about_ca_topic_score_gemma":0.008793457,"teacher_disagreement_score":0.009773025,"about_ca_system_score_codex":0.0004918531,"about_ca_system_score_gemma":0.000746435,"threshold_uncertainty_score":0.015188694},"labels":[],"label_agreement":null},{"id":"W4385491362","doi":"10.1016/j.jss.2023.111817","title":"On the impact of single and co-occurrent refactorings on quality attributes in android applications","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Android (operating system); Computer science; Software quality; Software; Software engineering; Java; Quality (philosophy); Software development; Programming language; Operating system","score_opus":0.07383912175224895,"score_gpt":0.3548067283088792,"score_spread":0.2809676065566303,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385491362","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99644995,0.0010105483,0.0008699942,0.00021318987,0.000053879252,0.00002162948,0.00026922266,0.00022923296,0.000882314],"genre_scores_gemma":[0.99581796,0.00026890772,0.0023108225,0.00007497196,0.000021527241,0.000011310829,0.0006872286,0.00009421983,0.00071310526],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9911521,0.001984851,0.0007781558,0.001308773,0.0040065222,0.000769582],"domain_scores_gemma":[0.741697,0.21656102,0.01047157,0.009370627,0.019402318,0.0024973685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062103965,0.0006382199,0.00050852523,0.002639321,0.0007744199,0.001383771,0.0011324175,0.0009883317,0.0016686281],"category_scores_gemma":[0.07551552,0.00033924473,0.0010390689,0.0017514837,0.00083037774,0.002739218,0.0012323058,0.0017296555,0.00035962518],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0057372143,0.002996248,0.5745918,0.0015261079,0.00143567,0.0025029194,0.0022921732,0.0521249,0.05358303,0.002035118,0.0056461953,0.29552862],"study_design_scores_gemma":[0.0002504137,0.0041327504,0.64548874,0.0004963254,0.0037313616,0.0021564644,0.004954564,0.29212797,0.036314677,0.003714511,0.00638791,0.0002443853],"about_ca_topic_score_codex":0.0148170795,"about_ca_topic_score_gemma":0.029788792,"teacher_disagreement_score":0.0148170795,"about_ca_system_score_codex":0.001146177,"about_ca_system_score_gemma":0.0020790743,"threshold_uncertainty_score":0.032844126},"labels":[],"label_agreement":null},{"id":"W4385519566","doi":"10.1613/jair.1.14394","title":"Program Synthesis with Best-First Bottom-Up Search","year":2023,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Search cost; Search algorithm; Beam search; Function (biology); Best-first search; String (physics); Theoretical computer science; Algorithm; Mathematics","score_opus":0.17891732954001874,"score_gpt":0.4215244089406741,"score_spread":0.24260707940065537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385519566","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032601073,0.00044039896,0.95771486,0.00017304024,0.00004189783,0.00013493003,0.00011736135,0.0029884023,0.005788141],"genre_scores_gemma":[0.34520823,0.0002525775,0.6488467,0.00028306525,0.000021671318,0.00035935736,0.00043954828,0.00075208285,0.0038367363],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99948406,0.00012150888,0.00003269644,0.00009722722,0.000191421,0.000073077084],"domain_scores_gemma":[0.9990062,0.00060478726,0.00006710285,0.00014758114,0.00014286088,0.000031541484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00061360025,0.0012338574,0.00088902604,0.0011834528,0.00049661833,0.0009818607,0.0011909452,0.0011097408,0.0043013687],"category_scores_gemma":[0.002520327,0.00058676564,0.0008889972,0.0008367196,0.0007649879,0.0012209115,0.0012154197,0.0011104859,0.00076905766],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020488116,0.00016745881,0.0010291826,0.00025493154,0.0000795517,0.00013078836,0.00012498612,0.67574185,0.014099319,0.023484703,0.0039319703,0.28075033],"study_design_scores_gemma":[0.000028795112,0.00006372988,0.00007902101,0.000015089956,0.00002351148,0.000028864615,0.000018070494,0.98528755,0.0031425238,0.0098305885,0.0014729591,0.000009242136],"about_ca_topic_score_codex":0.0029962754,"about_ca_topic_score_gemma":0.004430018,"teacher_disagreement_score":0.0043013687,"about_ca_system_score_codex":0.00093779335,"about_ca_system_score_gemma":0.001800865,"threshold_uncertainty_score":0.014389515},"labels":[],"label_agreement":null},{"id":"W4385572298","doi":"10.18653/v1/2023.findings-acl.280","title":"Debiasing should be Good and Bad: Measuring the Consistency of Debiasing Techniques in Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Debiasing; Computer science; Consistency (knowledge bases); Protocol (science); Reachability; Journaling file system; Artificial intelligence; Database; Theoretical computer science; Psychology","score_opus":0.08020725600439323,"score_gpt":0.3094917121458611,"score_spread":0.22928445614146786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572298","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.65057164,0.00031606757,0.33697483,0.0014523373,0.00008906075,0.0017851093,0.00051151897,0.0015274384,0.0067719193],"genre_scores_gemma":[0.8789274,0.00008211432,0.11840542,0.00026080103,0.000025250944,0.0009758923,0.00043369902,0.00024340342,0.0006460806],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.90823793,0.047357973,0.011833421,0.00610752,0.024763294,0.0016998735],"domain_scores_gemma":[0.44909346,0.35996348,0.055718042,0.08961815,0.042909082,0.0026977733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07359125,0.00073354284,0.00058130635,0.0035901056,0.0014828913,0.004211209,0.0021053425,0.0034371337,0.0012114062],"category_scores_gemma":[0.42249092,0.0008070455,0.00075155444,0.0021195333,0.004972462,0.007677687,0.0056448304,0.003165845,0.0003898851],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043074377,0.0027907263,0.3234978,0.0018850819,0.00091775885,0.00062247756,0.033349533,0.04213865,0.08066209,0.08659077,0.004768595,0.4184691],"study_design_scores_gemma":[0.00073872023,0.008335501,0.18699282,0.0011910186,0.0008533267,0.0014199784,0.015102494,0.3102699,0.2744917,0.1797164,0.019888422,0.0009997467],"about_ca_topic_score_codex":0.0012734405,"about_ca_topic_score_gemma":0.0013697388,"teacher_disagreement_score":0.07359125,"about_ca_system_score_codex":0.0022551827,"about_ca_system_score_gemma":0.0026569106,"threshold_uncertainty_score":0.38919234},"labels":[],"label_agreement":null},{"id":"W4385573986","doi":"10.18653/v1/2022.findings-emnlp.174","title":"CodeExp: Explanatory Code Document Generation","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Code (set theory); Task (project management); Baseline (sea); Code generation; Source code; Software; Source lines of code; Machine learning; Scale (ratio); Data mining; Artificial intelligence; Information retrieval; Programming language; Key (lock); Set (abstract data type)","score_opus":0.02828526137936499,"score_gpt":0.2698888864403468,"score_spread":0.24160362506098182,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573986","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06838025,0.0012914723,0.57668537,0.0019370699,0.00066687574,0.0012698745,0.04551535,0.28845468,0.015799101],"genre_scores_gemma":[0.22862236,0.00066403684,0.6062029,0.0010255793,0.000110221044,0.0013281491,0.13739876,0.009943427,0.014704544],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99894303,0.00033800394,0.00007238695,0.0003306393,0.00025180855,0.00006406606],"domain_scores_gemma":[0.99455655,0.002846582,0.00026151544,0.0012789278,0.00090712786,0.00014931924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014916316,0.0015403272,0.00039499864,0.0013977748,0.00033985634,0.001188935,0.0021443518,0.0012973864,0.009976814],"category_scores_gemma":[0.0113390805,0.00050389965,0.0009405275,0.0009138711,0.0004264163,0.002009033,0.0015640744,0.0017364399,0.005959145],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054636144,0.000467354,0.01160233,0.0014901186,0.00013124809,0.0005072963,0.0007509683,0.047268238,0.011701787,0.011031301,0.35109454,0.5634085],"study_design_scores_gemma":[0.00025045188,0.00028545552,0.0056632073,0.00025510072,0.00007514431,0.0005630279,0.00030773954,0.7564,0.038080562,0.019902756,0.17810385,0.00011277109],"about_ca_topic_score_codex":0.0047754147,"about_ca_topic_score_gemma":0.010039025,"teacher_disagreement_score":0.009976814,"about_ca_system_score_codex":0.0007907682,"about_ca_system_score_gemma":0.0013869426,"threshold_uncertainty_score":0.03337574},"labels":[],"label_agreement":null},{"id":"W4385607919","doi":"10.61187/ita.v1i1.19","title":"An Empirical Study of Adoption of ChatGPT for Bug Fixing among Professional Developers","year":2023,"lang":"en","type":"article","venue":"Innovation & Technology Advances","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Conversation; Expectancy theory; Computer science; Empirical research; Software development; Software; Key (lock); Knowledge management; Software engineering; Software peer review; Software construction; Computer security; Psychology","score_opus":0.03436966530522471,"score_gpt":0.378796272247145,"score_spread":0.3444266069419203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385607919","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9988464,0.000044506345,0.00042494465,0.0001022264,0.0000026738414,0.000033143268,0.00001065883,0.0000040844493,0.00053137005],"genre_scores_gemma":[0.99854505,0.00011769235,0.0008213485,0.000071979026,0.000004590258,0.000070850394,0.00002226347,0.0000044886783,0.00034176314],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9932459,0.003947004,0.00040305173,0.00054182566,0.0012516547,0.0006105991],"domain_scores_gemma":[0.8386814,0.11353535,0.01794448,0.0035552087,0.019598968,0.0066845226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011068632,0.00025301328,0.00030556007,0.00163549,0.0014326613,0.0015392079,0.00062825804,0.00081892754,0.0009763509],"category_scores_gemma":[0.079089984,0.0004901939,0.00021022517,0.0011348428,0.0011285751,0.0015348392,0.0010466395,0.0015117883,0.00017891773],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016870628,0.0011734697,0.54474676,0.00030453847,0.000031167165,0.0011320622,0.39058644,0.00016767027,0.0032141188,0.00087872264,0.0006692238,0.05692706],"study_design_scores_gemma":[0.00005230487,0.0013432496,0.6216941,0.00031738987,0.000055194294,0.0008970529,0.36392507,0.0029806697,0.0015942701,0.00041769838,0.0066535203,0.00006949898],"about_ca_topic_score_codex":0.0064643957,"about_ca_topic_score_gemma":0.010951801,"teacher_disagreement_score":0.011068632,"about_ca_system_score_codex":0.0016439783,"about_ca_system_score_gemma":0.0031534622,"threshold_uncertainty_score":0.058537185},"labels":[],"label_agreement":null},{"id":"W4385632434","doi":"10.1002/smr.2602","title":"Understanding the quality and evolution of Android app build systems","year":2023,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Australian Research Council; Monash University","keywords":"Android (operating system); Computer science; Plug-in; Executable; Scripting language; Software quality; Software engineering; Software; World Wide Web; Software development; Operating system","score_opus":0.0747303414186689,"score_gpt":0.32579833055526153,"score_spread":0.2510679891365926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385632434","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9936168,0.00028014474,0.0027344725,0.00028995101,0.0000054568245,0.000041528878,0.00010859278,0.00006052196,0.0028623638],"genre_scores_gemma":[0.9976878,0.0000956639,0.0016756329,0.000019508298,0.0000035097742,0.000014104412,0.00011809479,0.000027179705,0.00035849118],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99282795,0.0020147911,0.00064637227,0.0006707121,0.003334908,0.0005053734],"domain_scores_gemma":[0.88545805,0.050088145,0.027779985,0.007116809,0.027187934,0.002369058],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008469536,0.00025328045,0.00022104935,0.0043578735,0.0009186732,0.0037264302,0.00055042276,0.0006912506,0.0010921817],"category_scores_gemma":[0.08174053,0.0004556646,0.0002800251,0.0022595536,0.0015839301,0.0038145953,0.0021322588,0.0008935895,0.00025368165],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001273068,0.00010919182,0.89993495,0.00016990752,0.00005859451,0.0003670503,0.02019782,0.0025161968,0.0028769132,0.0026988802,0.0007196921,0.070223406],"study_design_scores_gemma":[0.000008672875,0.0001611434,0.96827567,0.00011830253,0.00004554922,0.0003797633,0.010673468,0.010771429,0.0021083504,0.0018405586,0.005571735,0.00004544526],"about_ca_topic_score_codex":0.009767591,"about_ca_topic_score_gemma":0.010448088,"teacher_disagreement_score":0.009767591,"about_ca_system_score_codex":0.0022795505,"about_ca_system_score_gemma":0.001233868,"threshold_uncertainty_score":0.0447917},"labels":[],"label_agreement":null},{"id":"W4385681195","doi":"10.1145/3610088","title":"SUMMIT: Scaffolding Open Source Software Issue Discussion Through Summarization","year":2023,"lang":"en","type":"article","venue":"Proceedings of the ACM on Human-Computer Interaction","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; McGill University","funders":"Alfred P. Sloan Foundation","keywords":"Automatic summarization; Summit; Computer science; World Wide Web; Set (abstract data type); Software; Data science; Construct (python library); Software engineering; Knowledge management; Information retrieval","score_opus":0.0643405934391388,"score_gpt":0.3516934090470921,"score_spread":0.28735281560795334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385681195","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1274607,0.0010751784,0.8130578,0.0024780452,0.0006817377,0.0032468755,0.0024025,0.033727776,0.015869416],"genre_scores_gemma":[0.26114425,0.00067956775,0.7134067,0.00072202215,0.000505875,0.0037411088,0.006411735,0.00382293,0.00956575],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98522687,0.009211373,0.0010790302,0.0018590182,0.0021915787,0.00043216502],"domain_scores_gemma":[0.87127537,0.09768244,0.007275425,0.010815399,0.009794655,0.0031565595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016256064,0.0026557066,0.0009869899,0.0066968086,0.0025169763,0.0056417454,0.0028157977,0.0017240806,0.009056034],"category_scores_gemma":[0.11909445,0.0008294501,0.0011550494,0.0027571649,0.0016022384,0.009982599,0.010807591,0.0020884373,0.004038133],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009504039,0.0009017094,0.010507308,0.004634784,0.00017468836,0.0008609576,0.11846794,0.004155852,0.034385342,0.013595697,0.04218173,0.7691837],"study_design_scores_gemma":[0.00079371413,0.0037413144,0.02731526,0.0045517054,0.00082129904,0.0014649095,0.08012926,0.13899893,0.06727063,0.10414613,0.56967235,0.0010945497],"about_ca_topic_score_codex":0.00092779327,"about_ca_topic_score_gemma":0.0021587845,"teacher_disagreement_score":0.016256064,"about_ca_system_score_codex":0.001079501,"about_ca_system_score_gemma":0.0021800976,"threshold_uncertainty_score":0.085971296},"labels":[],"label_agreement":null},{"id":"W4385768170","doi":"10.24963/ijcai.2023/328","title":"Can You Improve My Code? Optimizing Programs with Local Search","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; University of Alberta; Compute Canada","keywords":"Computer science; Code (set theory); Exploit; Source lines of code; Programming language; Programming paradigm; Line (geometry); Computer security; Set (abstract data type); Software; Mathematics","score_opus":0.02602334440744724,"score_gpt":0.272104756702353,"score_spread":0.24608141229490577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385768170","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051134292,0.00041651446,0.9355067,0.0006863892,0.000051449006,0.00009864902,0.000047945687,0.006222247,0.005835819],"genre_scores_gemma":[0.26386487,0.0003351983,0.72915095,0.0002848399,0.000038002818,0.00020777715,0.00013276245,0.001855481,0.0041302117],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985594,0.0005850657,0.0000684713,0.0002877904,0.0003820713,0.000117143776],"domain_scores_gemma":[0.996786,0.0019487598,0.0003081666,0.0005033909,0.00036703603,0.0000867034],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023910748,0.00084678235,0.00050359307,0.00048159692,0.0005097679,0.00080348545,0.0012687057,0.00069990987,0.0052886913],"category_scores_gemma":[0.010794151,0.00047276472,0.0005260056,0.00041962691,0.001199047,0.0020849442,0.0012597134,0.0014995921,0.0014270445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054919,0.00045757234,0.005025629,0.0007642624,0.0001250338,0.00027689346,0.0014566685,0.2401868,0.041695878,0.05130076,0.010898249,0.647263],"study_design_scores_gemma":[0.00016305859,0.0005459756,0.00086847757,0.00019105506,0.00009548886,0.00024801702,0.00036429105,0.8850898,0.038562413,0.046544507,0.027277693,0.000049154252],"about_ca_topic_score_codex":0.0010906408,"about_ca_topic_score_gemma":0.0021396321,"teacher_disagreement_score":0.0052886913,"about_ca_system_score_codex":0.00055193045,"about_ca_system_score_gemma":0.00094057585,"threshold_uncertainty_score":0.017692447},"labels":[],"label_agreement":null},{"id":"W4385848632","doi":"10.1145/3597503.3639154","title":"A User-centered Security Evaluation of Copilot","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Computer security; Code (set theory); Set (abstract data type)","score_opus":0.07139789347146061,"score_gpt":0.3556885183486882,"score_spread":0.28429062487722756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385848632","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9676967,0.00024797706,0.021179738,0.00024401036,0.000067311215,0.0010002234,0.0008384865,0.0038437007,0.004881906],"genre_scores_gemma":[0.92362237,0.00017100885,0.06598759,0.000406953,0.00003666743,0.0017443934,0.0024950823,0.0015526257,0.003983191],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99402654,0.0032878886,0.00041649977,0.00062840275,0.0012780573,0.00036249112],"domain_scores_gemma":[0.9302218,0.048837595,0.0030320822,0.0073549757,0.008450691,0.0021028875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072131264,0.0011824173,0.00065264705,0.0011559322,0.0005524231,0.0013707913,0.0013882713,0.0010776211,0.0049437145],"category_scores_gemma":[0.04571824,0.0003891905,0.0005172451,0.00048655286,0.00090421544,0.0015054495,0.0019510736,0.00132766,0.0013286986],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.02571083,0.01412127,0.11025038,0.008211432,0.00067959924,0.0023230957,0.030324968,0.014702768,0.224316,0.0046879184,0.040336497,0.5243352],"study_design_scores_gemma":[0.0041983672,0.05076068,0.32345805,0.001772138,0.00084798987,0.004884377,0.012029366,0.20799899,0.26811326,0.0072221933,0.1177565,0.00095806184],"about_ca_topic_score_codex":0.0009499309,"about_ca_topic_score_gemma":0.0016328426,"teacher_disagreement_score":0.0072131264,"about_ca_system_score_codex":0.0005414395,"about_ca_system_score_gemma":0.00064760406,"threshold_uncertainty_score":0.03814715},"labels":[],"label_agreement":null},{"id":"W4385965483","doi":"10.48550/arxiv.2308.08033","title":"Domain Adaptation for Code Model-based Unit Test Case Generation","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates","keywords":"Computer science; Leverage (statistics); Unit testing; Domain adaptation; Code coverage; Artificial intelligence; Machine learning; Transformer; Task (project management); Source code; Test data; Adaptation (eye); Test case; Data mining; Software; Software engineering; Programming language; Engineering","score_opus":0.23939016572429933,"score_gpt":0.2420723057480061,"score_spread":0.0026821400237067583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385965483","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18803018,0.0015455579,0.76186043,0.0005497088,0.00016459588,0.00040668348,0.002128442,0.03945634,0.005858027],"genre_scores_gemma":[0.69253975,0.00037220502,0.2941643,0.00057794835,0.000034490156,0.00045564995,0.0077571636,0.0014859986,0.002612527],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99851793,0.00053489103,0.00011509914,0.0004159104,0.0002911282,0.00012501521],"domain_scores_gemma":[0.99587554,0.0021088433,0.0003068138,0.00088140054,0.0007096667,0.00011777869],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001396034,0.0011765256,0.0005382195,0.0012531327,0.00018329614,0.0006574562,0.0017495404,0.00076352,0.0024170212],"category_scores_gemma":[0.008358624,0.00039110126,0.0009975515,0.0009925696,0.00052922365,0.0014351712,0.0010986782,0.001671959,0.0011754144],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003431004,0.00050813664,0.012027055,0.0004500298,0.00014849298,0.00048668918,0.00026501567,0.34081885,0.032210343,0.0038079747,0.016096136,0.5928382],"study_design_scores_gemma":[0.0000581355,0.000135332,0.0012018073,0.000028877503,0.00003667732,0.00018498098,0.000054116823,0.9712673,0.017872415,0.004266193,0.00487153,0.000022649767],"about_ca_topic_score_codex":0.0033504122,"about_ca_topic_score_gemma":0.004838611,"teacher_disagreement_score":0.0033504122,"about_ca_system_score_codex":0.0008882517,"about_ca_system_score_gemma":0.0012108237,"threshold_uncertainty_score":0.008085787},"labels":[],"label_agreement":null},{"id":"W4386158789","doi":"10.1109/csci58124.2022.00339","title":"Improving Quality of Software Requirements by Using a Triplet Structure","year":2022,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Software requirements specification; Software requirements; Software engineering; Requirements analysis; Software construction; Software development; Requirement; Verification and validation; Software system; Software quality; Requirements engineering; Software; Programming language; Engineering","score_opus":0.0487905860566675,"score_gpt":0.323466724958686,"score_spread":0.2746761389020185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386158789","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015131185,0.000081317914,0.9751789,0.00042835312,0.00008103165,0.00023363758,0.00034787587,0.0021154662,0.0064022257],"genre_scores_gemma":[0.09050805,0.00013421765,0.9054201,0.000117343814,0.00003557089,0.00027434743,0.000767447,0.0005522784,0.0021906714],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9930362,0.002279972,0.0008446938,0.00094138226,0.0026398774,0.00025795813],"domain_scores_gemma":[0.9842278,0.006677188,0.001647926,0.0039671278,0.003161902,0.00031810848],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049766116,0.0005522933,0.0009084079,0.00376953,0.0016100492,0.0025358724,0.0011806047,0.0013484065,0.004794653],"category_scores_gemma":[0.020297468,0.00065882836,0.0016614991,0.003973806,0.0018097148,0.005365115,0.0020754326,0.0023538957,0.0014866379],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040203295,0.00046064172,0.005219752,0.00079979654,0.00009066966,0.0013815212,0.004396625,0.043852497,0.037201162,0.5006883,0.011753667,0.3937533],"study_design_scores_gemma":[0.00018913837,0.00066458527,0.0022678378,0.0004126793,0.00014882824,0.0012666553,0.00082664815,0.44690996,0.041409425,0.39632767,0.10942136,0.00015522928],"about_ca_topic_score_codex":0.0026414834,"about_ca_topic_score_gemma":0.0036564625,"teacher_disagreement_score":0.0049766116,"about_ca_system_score_codex":0.0015586717,"about_ca_system_score_gemma":0.0025590288,"threshold_uncertainty_score":0.026319146},"labels":[],"label_agreement":null},{"id":"W4386212357","doi":"10.1109/tse.2023.3307243","title":"ADPTriage: Approximate Dynamic Programming for Bug Triage","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Toronto Metropolitan University","funders":"","keywords":"Computer science; Triage; Markov decision process; Process (computing); Software bug; Software regression; Task (project management); Software; Pipeline (software); Software engineering; Markov process; Software development; Programming language; Software quality; Systems engineering","score_opus":0.019948979039290816,"score_gpt":0.2724453377257305,"score_spread":0.25249635868643966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386212357","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0073207873,0.00044640462,0.9883762,0.000390636,0.00008573936,0.00009397589,0.00015545837,0.00046018467,0.0026707042],"genre_scores_gemma":[0.55207026,0.00081823004,0.43921295,0.0006480045,0.00015679646,0.0006833116,0.00068432145,0.00038839743,0.005337838],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99822646,0.00071486953,0.000089521905,0.00036009747,0.00035258775,0.0002564558],"domain_scores_gemma":[0.99246055,0.0060836836,0.0004380298,0.00025402717,0.00046873244,0.00029504814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032645494,0.0017516279,0.0025929702,0.00094496383,0.0008078345,0.0018758712,0.0022433184,0.0020608294,0.0052596545],"category_scores_gemma":[0.010986083,0.0013019156,0.0013389329,0.0013988061,0.001114726,0.0016451816,0.002254685,0.0035471069,0.000652707],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058081616,0.000048138452,0.0003415416,0.000086630374,0.00003102455,0.000034568326,0.000036212066,0.9724663,0.00017258551,0.0075316182,0.0013441826,0.01784909],"study_design_scores_gemma":[0.0000074298578,0.000013004884,0.000022373195,0.000005385126,0.0000038232974,0.0000049844043,0.0000048710303,0.9955852,0.00003802464,0.0040328945,0.00027937375,0.0000025328939],"about_ca_topic_score_codex":0.012092539,"about_ca_topic_score_gemma":0.010855676,"teacher_disagreement_score":0.012092539,"about_ca_system_score_codex":0.0018333024,"about_ca_system_score_gemma":0.0040429477,"threshold_uncertainty_score":0.024044275},"labels":[],"label_agreement":null},{"id":"W4386212466","doi":"10.1109/siu59756.2023.10223806","title":"Transformer-Based Bug/Feature Classification","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Support vector machine; Artificial intelligence; Convolutional neural network; Software bug; Transformer; Software; Feature vector; Artificial neural network; Machine learning; Feature (linguistics); Pattern recognition (psychology); Data mining; Engineering; Programming language","score_opus":0.03799938041550322,"score_gpt":0.291227708601163,"score_spread":0.25322832818565977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386212466","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7242446,0.0011011455,0.23451646,0.00038490572,0.00022826003,0.00031681097,0.006957449,0.02614084,0.0061094495],"genre_scores_gemma":[0.9403384,0.0002194214,0.046256077,0.00005332652,0.000029104378,0.000065786946,0.010356251,0.00016226775,0.0025193943],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99892104,0.00011802761,0.000118296484,0.00032772476,0.00038426882,0.00013071233],"domain_scores_gemma":[0.9978811,0.00043325697,0.00038330324,0.00041928302,0.00079453253,0.00008850031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009902716,0.0009355885,0.00069847034,0.0028257295,0.00020103846,0.0006993362,0.0009261673,0.00040348616,0.0013845697],"category_scores_gemma":[0.004151971,0.00016454075,0.00068764173,0.0016416305,0.00026641157,0.0013940895,0.000761225,0.00058052037,0.0012206167],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065303955,0.0006639073,0.09098498,0.00034361167,0.00017853969,0.00038358654,0.00017352038,0.038370665,0.02273197,0.0011957913,0.016295979,0.8280245],"study_design_scores_gemma":[0.000040233506,0.00035549334,0.04941524,0.000025203622,0.000116286144,0.0006086036,0.00013243072,0.9147378,0.027673457,0.0020215213,0.004824716,0.00004891393],"about_ca_topic_score_codex":0.007907662,"about_ca_topic_score_gemma":0.008649238,"teacher_disagreement_score":0.007907662,"about_ca_system_score_codex":0.00076855253,"about_ca_system_score_gemma":0.0008117793,"threshold_uncertainty_score":0.015723288},"labels":[],"label_agreement":null},{"id":"W4386213535","doi":"10.1002/smr.2609","title":"Image‐based communication on social coding platforms","year":2023,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software; Coding (social sciences); World Wide Web; Variety (cybernetics); Social media; Image (mathematics); Multimedia; Data science; Information retrieval; Software engineering; Artificial intelligence; Programming language","score_opus":0.036699246919948886,"score_gpt":0.32083804380023945,"score_spread":0.28413879688029053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386213535","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97684866,0.00031447643,0.0066744974,0.00059925433,0.00003913034,0.00015071903,0.0016146856,0.00016036049,0.013598177],"genre_scores_gemma":[0.9935656,0.00014096157,0.004087875,0.00006153393,0.000051300583,0.000071819944,0.0006414078,0.00004206266,0.0013374597],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9966935,0.0014425076,0.0001796234,0.0003814974,0.0010238562,0.00027907407],"domain_scores_gemma":[0.9454584,0.035128202,0.01096755,0.0024483118,0.0050197584,0.0009777442],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0023040914,0.0003801733,0.0002805007,0.0071860547,0.001054272,0.002641109,0.0005840289,0.0007357632,0.0038261372],"category_scores_gemma":[0.03528735,0.00023076883,0.000246836,0.003911419,0.0011384587,0.005383822,0.0019575655,0.0006229289,0.00073658855],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060267205,0.0003360393,0.6395119,0.0018179821,0.00018318658,0.0020631421,0.06293819,0.002194539,0.014217886,0.009610757,0.00884843,0.2576753],"study_design_scores_gemma":[0.000033477383,0.00044784063,0.79328126,0.0010607518,0.00020580206,0.0023763508,0.065551326,0.037946235,0.01354297,0.012361422,0.07292616,0.00026639993],"about_ca_topic_score_codex":0.0048291306,"about_ca_topic_score_gemma":0.0061786906,"teacher_disagreement_score":0.9973589,"about_ca_system_score_codex":0.0011018169,"about_ca_system_score_gemma":0.000620187,"threshold_uncertainty_score":0.01279974},"labels":[],"label_agreement":null},{"id":"W4386301778","doi":"10.48550/arxiv.2308.13963","title":"GPTCloneBench: A comprehensive benchmark of semantic clones and cross-language clones using GPT-3 model and SemanticCloneBench","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Benchmark (surveying); clone (Java method); Programming language; Java; Artificial intelligence; Natural language processing; Python (programming language); Machine learning; Software engineering","score_opus":0.1163536748243747,"score_gpt":0.2545266693069382,"score_spread":0.13817299448256348,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386301778","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5991944,0.0054646255,0.103757575,0.0020791623,0.0009680794,0.0013892426,0.123957135,0.13634755,0.026842257],"genre_scores_gemma":[0.3701894,0.001221641,0.15269856,0.0011652083,0.000085633714,0.0012923962,0.45867917,0.0089351125,0.0057328544],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9933234,0.0011699493,0.00063905754,0.001881023,0.002499145,0.00048739935],"domain_scores_gemma":[0.9895978,0.0037971966,0.000770203,0.002711115,0.0025220555,0.0006015686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036440196,0.0030740967,0.000949906,0.0043245405,0.0011609248,0.0021770461,0.0056596454,0.0024516678,0.0022761554],"category_scores_gemma":[0.018654805,0.0007848985,0.0022681085,0.0052096,0.0016186266,0.0040873806,0.0033253394,0.0025929615,0.0019376932],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023949384,0.0022049511,0.0892474,0.005459394,0.0010776647,0.002873801,0.0015119371,0.2230483,0.031142471,0.014686675,0.34217757,0.28417492],"study_design_scores_gemma":[0.00069494464,0.0017172361,0.03341089,0.00044339788,0.00035360234,0.0019129177,0.0010846267,0.74703056,0.053492717,0.015061487,0.14454862,0.00024895623],"about_ca_topic_score_codex":0.016404817,"about_ca_topic_score_gemma":0.020367859,"teacher_disagreement_score":0.016404817,"about_ca_system_score_codex":0.0023408772,"about_ca_system_score_gemma":0.003133784,"threshold_uncertainty_score":0.03261864},"labels":[],"label_agreement":null},{"id":"W4386365732","doi":"10.18293/seke2023-176","title":"An Optimal Spacing Approach for Sampling Small-sized Datasets","year":2023,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Machine learning; Computer science; Artificial intelligence; Ranking (information retrieval); Deep learning; Context (archaeology); Focus (optics); Sampling (signal processing); Data mining","score_opus":0.052366070853110874,"score_gpt":0.29275279852208924,"score_spread":0.24038672766897837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386365732","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03198603,0.00033885622,0.96532136,0.00031961603,0.00006105717,0.00030538868,0.00023815899,0.00054845674,0.00088100607],"genre_scores_gemma":[0.49297917,0.00028014896,0.5015936,0.00042571372,0.00015234864,0.0013550944,0.0019000869,0.00014927448,0.0011645722],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9861173,0.008051855,0.0009261394,0.0025761917,0.0019012972,0.0004272219],"domain_scores_gemma":[0.9448414,0.042062312,0.003151284,0.005673201,0.0034311037,0.00084073073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025674792,0.0009990537,0.0020998649,0.003130507,0.0010540818,0.0017472568,0.003265017,0.0022067744,0.0021187298],"category_scores_gemma":[0.07209156,0.000727026,0.0016968925,0.0029583292,0.0018493702,0.0031302955,0.0028731248,0.0033201613,0.00045629896],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012814936,0.0010081627,0.027214222,0.00065413414,0.0005266688,0.00045499668,0.0008616892,0.57348555,0.0042689536,0.062184658,0.007037385,0.32102212],"study_design_scores_gemma":[0.000104446706,0.00029137946,0.0020091163,0.00007148771,0.000036705944,0.000080937396,0.0001352377,0.9529084,0.001387421,0.04167972,0.0012712836,0.000023777106],"about_ca_topic_score_codex":0.0021774182,"about_ca_topic_score_gemma":0.0022019143,"teacher_disagreement_score":0.025674792,"about_ca_system_score_codex":0.0015655555,"about_ca_system_score_gemma":0.0021569128,"threshold_uncertainty_score":0.1357829},"labels":[],"label_agreement":null},{"id":"W4386564509","doi":"10.1177/01454455231195825","title":"Further Progress Toward Automating Functional Analysis Interpretation","year":2023,"lang":"en","type":"article","venue":"Behavior Modification","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"","keywords":"Psychology; Interpretation (philosophy); Functional analysis; Cognitive psychology; Cognitive science; Computer science; Programming language; Chemistry","score_opus":0.07077031815588423,"score_gpt":0.3370757885863228,"score_spread":0.26630547043043856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386564509","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020584324,0.00118451,0.9490229,0.004524998,0.00036165942,0.0012625706,0.0006689385,0.016186379,0.006203692],"genre_scores_gemma":[0.026038568,0.00044247517,0.9691037,0.0004164413,0.0001133,0.00034037587,0.00083285756,0.0010915262,0.00162081],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.936621,0.040821236,0.0043191724,0.0056551523,0.011665637,0.0009178368],"domain_scores_gemma":[0.6152672,0.18296741,0.014539843,0.09689197,0.08784385,0.0024896597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08284421,0.0033815792,0.002115199,0.0068616434,0.0010464735,0.0070847156,0.007614062,0.002049174,0.013352925],"category_scores_gemma":[0.17385462,0.0013695492,0.002438871,0.0040700235,0.0028615245,0.00787389,0.003365448,0.004356419,0.007826342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003233669,0.00075221964,0.008864167,0.001866736,0.00027092668,0.00013572564,0.0050005326,0.004196248,0.023264963,0.008213858,0.022677848,0.92443335],"study_design_scores_gemma":[0.0007434796,0.0019912526,0.06664879,0.006907331,0.00064742146,0.0020153055,0.0073853885,0.24476604,0.11623433,0.075494535,0.47579455,0.0013716242],"about_ca_topic_score_codex":0.015166412,"about_ca_topic_score_gemma":0.010468608,"teacher_disagreement_score":0.08284421,"about_ca_system_score_codex":0.0026099142,"about_ca_system_score_gemma":0.007468811,"threshold_uncertainty_score":0.43812728},"labels":[],"label_agreement":null},{"id":"W4386601484","doi":"10.3390/electronics12183805","title":"Software Requirement Risk Prediction Using Enhanced Fuzzy Induction Models","year":2023,"lang":"en","type":"article","venue":"Electronics","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Universiti Teknologi Petronas; Sexually Transmitted Infection Research Foundation","keywords":"Software development; Computer science; Risk analysis (engineering); Verification and validation; Software development process; Goal-Driven Software Development Process; Software construction; Software system; Software; Reliability engineering; Package development process; Software engineering; Engineering; Operations management","score_opus":0.03628965410702738,"score_gpt":0.28121169712756866,"score_spread":0.24492204302054127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386601484","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.071453586,0.00018956937,0.9257706,0.00011206289,0.00001713676,0.0000570885,0.000088338056,0.0002056932,0.0021060212],"genre_scores_gemma":[0.90669805,0.00019533098,0.0912671,0.000042324744,0.000017070079,0.000109679895,0.00017496455,0.000014693086,0.0014808115],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99933296,0.00018190117,0.00005005055,0.0001406066,0.00023556717,0.000058911228],"domain_scores_gemma":[0.99868137,0.0007806834,0.00018437601,0.000051991705,0.00027471158,0.000026893471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011581013,0.0005218079,0.00061515404,0.00093656295,0.00027826778,0.0007025513,0.0009472065,0.0005748574,0.0008445844],"category_scores_gemma":[0.0029647802,0.00023601991,0.0007808242,0.0005134935,0.00025032467,0.0007181215,0.00037681998,0.00074136496,0.00018287385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006901966,0.00005413166,0.0023705566,0.000040160532,0.000034195153,0.00007208979,0.00007153943,0.9451189,0.0016001716,0.004146644,0.00034568366,0.046076886],"study_design_scores_gemma":[0.0000012239452,0.0000067850983,0.00015181486,0.0000022855713,0.0000027261235,0.000004318378,0.0000024766946,0.9989562,0.00016172572,0.00066311844,0.000045246994,0.0000019483286],"about_ca_topic_score_codex":0.0072848243,"about_ca_topic_score_gemma":0.0044380506,"teacher_disagreement_score":0.0072848243,"about_ca_system_score_codex":0.00068791735,"about_ca_system_score_gemma":0.0007036133,"threshold_uncertainty_score":0.014484823},"labels":[],"label_agreement":null},{"id":"W4386643890","doi":"10.1007/978-3-030-02391-1_1","title":"Software Life Cycle","year":2022,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"British Columbia Institute of Technology","funders":"","keywords":"CLARITY; Computer science; Software development process; Objectivity (philosophy); Software engineering; Software development; Software; Variety (cybernetics); Structuring; Stakeholder; Software deployment; World Wide Web; Artificial intelligence; Management","score_opus":0.02114967138053906,"score_gpt":0.24239370981007166,"score_spread":0.2212440384295326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386643890","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022126054,0.0028586818,0.054174952,0.0008300092,0.00025152045,0.00009715379,0.00049799413,0.00078166253,0.9382953],"genre_scores_gemma":[0.03416505,0.0054833163,0.018202202,0.00043551894,0.000107541535,0.000120933175,0.0011133554,0.000708956,0.93966305],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997011,0.000041837713,0.000010333808,0.00004359035,0.00018095461,0.000022212576],"domain_scores_gemma":[0.9996784,0.00008670292,0.000015344456,0.00007064457,0.00012670363,0.000022207476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00025831826,0.0007031373,0.00028302587,0.0013479661,0.0005537212,0.0022338235,0.000719832,0.0004829105,0.063693546],"category_scores_gemma":[0.0010702412,0.00038076864,0.00031847833,0.0013513307,0.00049818505,0.0023263863,0.0007125808,0.00097461004,0.029308138],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020020369,0.00005240369,0.00014335277,0.00022071911,0.0000069428547,0.00006311749,0.00028667247,0.004105062,0.0027646956,0.41592926,0.12354457,0.45286322],"study_design_scores_gemma":[0.0000030647495,0.00002536519,0.00016747175,0.00010898446,0.0000044841872,0.00011259605,0.000064453576,0.0027675524,0.0012341316,0.1089944,0.8865101,0.000007428767],"about_ca_topic_score_codex":0.0027616143,"about_ca_topic_score_gemma":0.0028848778,"teacher_disagreement_score":0.063693546,"about_ca_system_score_codex":0.0016955555,"about_ca_system_score_gemma":0.0015653565,"threshold_uncertainty_score":0.21307617},"labels":[],"label_agreement":null},{"id":"W4386715810","doi":"10.1007/s10664-023-10373-0","title":"Investigating developers’ perception on software testability and its effects","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Dalhousie University","funders":"","keywords":"Testability; Code smell; Software engineering; Computer science; Empirical research; Software quality; Software development; Software; Engineering; Reliability engineering; Programming language","score_opus":0.03623200916931937,"score_gpt":0.2923139179714733,"score_spread":0.256081908802154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386715810","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9978999,0.000047471774,0.0003645109,0.00008156583,0.0000035414564,0.0000052584687,0.000016245893,0.0000074017175,0.0015740468],"genre_scores_gemma":[0.99954396,0.000023436918,0.00019042355,0.00001556655,0.000002894033,0.0000034689535,0.000021671678,0.0000037363955,0.00019487648],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99600947,0.0018055593,0.0003344837,0.00030078168,0.0012772455,0.00027245443],"domain_scores_gemma":[0.7585537,0.19514947,0.023039693,0.0048852884,0.013427385,0.0049444474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006018847,0.00022811462,0.0001781391,0.00089085894,0.00026088394,0.0011232396,0.00031641367,0.0005974649,0.002964831],"category_scores_gemma":[0.10023458,0.00019620733,0.00028285215,0.0005612292,0.0005712501,0.0011524975,0.00077338493,0.00078457623,0.00021723019],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071028597,0.00069474557,0.93723565,0.0001624863,0.00009467489,0.00030570818,0.018544534,0.00049754104,0.00915129,0.00048212276,0.00033416055,0.031786766],"study_design_scores_gemma":[0.000025039522,0.0005424721,0.9885076,0.000038365473,0.00005542793,0.00015934216,0.0066933045,0.001353645,0.0016491835,0.0002883528,0.00066646124,0.00002077711],"about_ca_topic_score_codex":0.0037080683,"about_ca_topic_score_gemma":0.0047599524,"teacher_disagreement_score":0.006018847,"about_ca_system_score_codex":0.00058875495,"about_ca_system_score_gemma":0.00064226484,"threshold_uncertainty_score":0.031831086},"labels":[],"label_agreement":null},{"id":"W4386781823","doi":"10.1007/s10664-023-10347-2","title":"A study of documentation for software architecture","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Software documentation; Computer science; Internal documentation; Context (archaeology); World Wide Web; Software engineering; Software; Software system; Information retrieval; Programming language; Software construction","score_opus":0.0365751899518581,"score_gpt":0.329109199373997,"score_spread":0.2925340094221389,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386781823","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94915724,0.0011259953,0.009754143,0.0018739762,0.00002414026,0.000036760597,0.000055366057,0.00007344155,0.037898943],"genre_scores_gemma":[0.99092126,0.00047830862,0.0038857593,0.00008833377,0.000013489673,0.000016901811,0.000058839832,0.00003309564,0.004504045],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.99741864,0.0016070567,0.000117437,0.00011325624,0.00058014435,0.00016342825],"domain_scores_gemma":[0.92338234,0.06072633,0.00497383,0.0047759106,0.0049689324,0.0011726525],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0035713625,0.00023075256,0.00022123342,0.002232267,0.0024750682,0.0023407694,0.00073453336,0.0010417814,0.0033631779],"category_scores_gemma":[0.060008924,0.00033284153,0.00027852497,0.003920404,0.002632151,0.0049281195,0.001157751,0.0018654552,0.00032295243],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002987343,0.0015650856,0.20147085,0.0005599684,0.000051397456,0.001681837,0.08727778,0.0061358586,0.0040943916,0.43517914,0.007519256,0.25416574],"study_design_scores_gemma":[0.00022332859,0.001811058,0.34899426,0.0018650026,0.00018178973,0.005037546,0.10879419,0.056143153,0.010119697,0.30524445,0.161424,0.00016153062],"about_ca_topic_score_codex":0.008598512,"about_ca_topic_score_gemma":0.011250468,"teacher_disagreement_score":0.9964286,"about_ca_system_score_codex":0.0024824527,"about_ca_system_score_gemma":0.0039718268,"threshold_uncertainty_score":0.0188874},"labels":[],"label_agreement":null},{"id":"W4386825342","doi":"10.1109/tse.2023.3313875","title":"A Grounded Theory of Cross-Community SECOs: Feedback Diversity Versus Synchronization","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Grounded theory; Computer science; Synchronization (alternating current); Diversity (politics); Terabyte; Upstream (networking); Data science; World Wide Web; Operating system; Qualitative research; Computer network; Sociology","score_opus":0.039422114663781684,"score_gpt":0.27669226271369296,"score_spread":0.23727014804991128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386825342","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17134169,0.0006244581,0.7493337,0.014226176,0.00013803136,0.0022445282,0.00048156182,0.00017265904,0.06143719],"genre_scores_gemma":[0.86428636,0.0002792364,0.13191652,0.0007054314,0.000029041821,0.0017489049,0.00018537563,0.000034552013,0.0008145567],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9703946,0.021707471,0.0011837371,0.0027696083,0.0032248888,0.0007196949],"domain_scores_gemma":[0.9379835,0.047927644,0.0039577098,0.0046978514,0.00397342,0.0014598309],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030055461,0.0010425036,0.000863821,0.0057804463,0.0047545996,0.006673624,0.0028612756,0.0024294804,0.0022762138],"category_scores_gemma":[0.037267424,0.0008789523,0.0010693744,0.0041541797,0.0315471,0.011852937,0.0049065254,0.0032494983,0.00030092348],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000057966532,0.00016302294,0.008372629,0.0004904235,0.000049677383,0.00020410758,0.1029215,0.0032709981,0.00089340156,0.8567798,0.0008783828,0.025918063],"study_design_scores_gemma":[0.00014770955,0.000152517,0.0049031028,0.0009029751,0.000054369677,0.00020520958,0.06579667,0.0203771,0.0009954888,0.8819536,0.024445688,0.00006563278],"about_ca_topic_score_codex":0.006172457,"about_ca_topic_score_gemma":0.005535624,"teacher_disagreement_score":0.030055461,"about_ca_system_score_codex":0.012740826,"about_ca_system_score_gemma":0.008677708,"threshold_uncertainty_score":0.15895039},"labels":[],"label_agreement":null},{"id":"W4386830495","doi":"10.1145/3624740","title":"<i>LoGenText-Plus</i> : Improving Neural Machine Translation Based Logging Texts Generation with Syntactic Templates","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Waterloo","funders":"","keywords":"Computer science; Logging; Source code; Code (set theory); Context (archaeology); Template; Natural language processing; Artificial intelligence; Database; Programming language; Set (abstract data type)","score_opus":0.09871842548096321,"score_gpt":0.3114374323806604,"score_spread":0.2127190068996972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386830495","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06686246,0.0011915951,0.7669296,0.0019524817,0.0011298975,0.0006817658,0.005412722,0.14091447,0.014924954],"genre_scores_gemma":[0.19076154,0.00052544713,0.7572302,0.0013230776,0.00026731478,0.0005903712,0.02183441,0.0060637044,0.021403858],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987048,0.00039604882,0.00011322591,0.0003736582,0.00033190614,0.000080375234],"domain_scores_gemma":[0.9963819,0.0015133884,0.0001780248,0.0009165876,0.0008940443,0.00011601343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011888337,0.002038308,0.00073983136,0.001419848,0.00072149094,0.0016993895,0.0019559555,0.0015359246,0.009202589],"category_scores_gemma":[0.007228444,0.00054064766,0.0010393952,0.0010364604,0.0007756677,0.0033640552,0.0017889547,0.0021751567,0.0071123783],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044118922,0.0005734803,0.0017354698,0.0009358338,0.00014716099,0.00076550175,0.00060973095,0.030376699,0.042102322,0.0063924,0.11898388,0.7969364],"study_design_scores_gemma":[0.00022475966,0.00033205844,0.0016425181,0.00010324437,0.00012479327,0.00061510515,0.00028168576,0.8262599,0.0904217,0.008872005,0.070989765,0.00013244126],"about_ca_topic_score_codex":0.0061825044,"about_ca_topic_score_gemma":0.010846771,"teacher_disagreement_score":0.009202589,"about_ca_system_score_codex":0.0007677202,"about_ca_system_score_gemma":0.0021435786,"threshold_uncertainty_score":0.03078574},"labels":[],"label_agreement":null},{"id":"W4386869826","doi":"10.2139/ssrn.4568851","title":"Cognitive Biases in Natural Language: Automatically Detecting, Differentiating, and Measuring Bias in Text","year":2023,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Natural (archaeology); Computer science; Natural language processing; Cognitive bias; Cognition; Cognitive psychology; Natural language; Natural language understanding; Artificial intelligence; Psychology; Linguistics; Philosophy; Geography","score_opus":0.04473239820399759,"score_gpt":0.3088940874624128,"score_spread":0.2641616892584152,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386869826","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8733392,0.0016767357,0.11443434,0.0006322828,0.00020840035,0.0003440617,0.001718173,0.0029112855,0.0047354572],"genre_scores_gemma":[0.9277915,0.00038787627,0.068762876,0.00024968747,0.0001942292,0.000182279,0.0013750838,0.00022442543,0.00083209557],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959116,0.0012298011,0.00045464904,0.001046059,0.0011472346,0.00021064653],"domain_scores_gemma":[0.9372093,0.048388287,0.0066721626,0.0024539079,0.0043127453,0.0009636021],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041588554,0.0006535956,0.00087312097,0.0036775966,0.00040730165,0.0035360039,0.0008373588,0.0013155112,0.0020339456],"category_scores_gemma":[0.057188775,0.00035655792,0.0004834282,0.0020133178,0.0006982773,0.004357826,0.0014832194,0.0012033106,0.0014209842],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002230475,0.0006494718,0.13122278,0.0014988962,0.00035275292,0.00035121752,0.0039852527,0.0021382787,0.124617815,0.0030104853,0.0063367975,0.72360575],"study_design_scores_gemma":[0.00055329077,0.0019064489,0.46012405,0.0005563846,0.0012845993,0.0027457813,0.005260943,0.28190878,0.16550082,0.06563654,0.013994388,0.0005280776],"about_ca_topic_score_codex":0.0013650539,"about_ca_topic_score_gemma":0.0016423645,"teacher_disagreement_score":0.0041588554,"about_ca_system_score_codex":0.00052120513,"about_ca_system_score_gemma":0.0008344484,"threshold_uncertainty_score":0.021994412},"labels":[],"label_agreement":null},{"id":"W4386977908","doi":"10.1145/3597503.3623347","title":"BOMs Away! Inside the Minds of Stakeholders: A Comprehensive Study of Bills of Materials for Software Systems","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"National Science Foundation","keywords":"Stakeholder; Computer science; Knowledge management; Domain (mathematical analysis); Software; Software development; Engineering management; Data science; World Wide Web; Engineering; Public relations; Political science","score_opus":0.12032324391725269,"score_gpt":0.32092237571344395,"score_spread":0.20059913179619127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386977908","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9951126,0.00015275693,0.0007791164,0.001962296,0.0000062120544,0.00004721659,0.000020613377,0.0000071659224,0.001912043],"genre_scores_gemma":[0.99658066,0.00035521245,0.0013309288,0.0007554267,0.0000056939966,0.000076549346,0.00003473778,0.000020947873,0.00083996914],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9904389,0.0064219553,0.00035621924,0.00043306896,0.0011754937,0.0011744596],"domain_scores_gemma":[0.97120124,0.01951593,0.003344189,0.0007788975,0.0026810954,0.0024785616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017665196,0.00035065398,0.0004793356,0.0019834586,0.0065339757,0.0032576958,0.001121772,0.002262836,0.001636518],"category_scores_gemma":[0.03541469,0.0007653845,0.00030376733,0.0020997429,0.0051310966,0.009577249,0.004382322,0.0032213298,0.00030535882],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015927062,0.00011136301,0.015407838,0.000104358885,0.000003976343,0.00043310583,0.9743456,0.00003775579,0.0005648813,0.0011887449,0.00047345104,0.007312896],"study_design_scores_gemma":[0.0000030722565,0.000093142415,0.012817714,0.00011928764,0.0000038131468,0.00013985056,0.98049265,0.00015928509,0.00016014383,0.00047799628,0.00552127,0.000011777961],"about_ca_topic_score_codex":0.013429922,"about_ca_topic_score_gemma":0.03429716,"teacher_disagreement_score":0.017665196,"about_ca_system_score_codex":0.004232675,"about_ca_system_score_gemma":0.0074396497,"threshold_uncertainty_score":0.093423605},"labels":[],"label_agreement":null},{"id":"W4386982649","doi":"10.1007/s10664-023-10380-1","title":"Is GitHub’s Copilot as bad as humans at introducing vulnerabilities in code?","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":101,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Vulnerability (computing); Code (set theory); Process (computing); Computer security; Perspective (graphical); Secure coding; Software; Software engineering; Software security assurance; Artificial intelligence; Operating system; Information security; Programming language","score_opus":0.03474644801521013,"score_gpt":0.3154596742478399,"score_spread":0.28071322623262973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386982649","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6374408,0.0043500513,0.044341397,0.17143948,0.0025930752,0.00015216068,0.0013690192,0.010164866,0.12814924],"genre_scores_gemma":[0.95042783,0.0011409206,0.018258484,0.01746825,0.000322523,0.00005797809,0.00081321405,0.0023744698,0.009136404],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9867927,0.0056361733,0.00037819182,0.0013278206,0.004550868,0.0013143821],"domain_scores_gemma":[0.9246243,0.035595886,0.008131279,0.016111314,0.011203392,0.004333802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012840184,0.0007556767,0.00058057706,0.0021953646,0.0019537748,0.004201843,0.00148113,0.0035679808,0.0075339377],"category_scores_gemma":[0.11352643,0.0005736612,0.00046579458,0.0016228828,0.005434384,0.00982682,0.0030052043,0.0032896413,0.0031157413],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018738054,0.0006240356,0.2007792,0.0009830421,0.0005266014,0.0011321205,0.0146606015,0.0045946348,0.007666979,0.07533218,0.2782953,0.41353157],"study_design_scores_gemma":[0.00045318308,0.0016833603,0.19890307,0.0025186313,0.00073039427,0.0057502743,0.035737615,0.031470668,0.025863752,0.21326347,0.48300874,0.0006168806],"about_ca_topic_score_codex":0.012724229,"about_ca_topic_score_gemma":0.018767169,"teacher_disagreement_score":0.012840184,"about_ca_system_score_codex":0.001564931,"about_ca_system_score_gemma":0.0043486687,"threshold_uncertainty_score":0.0679062},"labels":[],"label_agreement":null},{"id":"W4387103442","doi":"10.21203/rs.3.rs-3369458/v1","title":"A Machine Learning Approach to Determine the Semantic Versioning Type of npm Packages Releases","year":2023,"lang":"en","type":"preprint","venue":"Research Square","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Software versioning; Computer science; Machine learning; Artificial intelligence; Random forest; Natural language processing; Software; Programming language","score_opus":0.12337096386953493,"score_gpt":0.3766146636647912,"score_spread":0.25324369979525624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387103442","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20069195,0.00063611043,0.7796771,0.000877355,0.00015960551,0.00021227052,0.0037366059,0.0074583213,0.0065507563],"genre_scores_gemma":[0.6702865,0.0001744776,0.3218076,0.00013515077,0.00012317194,0.00011650091,0.0045159496,0.00026526678,0.002575423],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977393,0.00029995083,0.0002622065,0.00090398843,0.0005993856,0.0001951754],"domain_scores_gemma":[0.9902264,0.0049131378,0.0013290208,0.0012775286,0.0020003894,0.00025347588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019805904,0.00075767154,0.00080105814,0.0046065277,0.00095610763,0.0027127992,0.0019613442,0.0014973361,0.002420224],"category_scores_gemma":[0.013202517,0.00053459016,0.0013722138,0.0023686816,0.00072327774,0.003550678,0.00092109526,0.0023671202,0.0017117294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008079978,0.0010301629,0.08075369,0.00031195182,0.00019094019,0.0003271669,0.00036051808,0.104810774,0.018820316,0.020934628,0.01255124,0.75910074],"study_design_scores_gemma":[0.000017051449,0.00004601008,0.00453218,0.000024508465,0.00003110912,0.00011452721,0.00005061987,0.97473043,0.0042908075,0.014869384,0.001275747,0.000017582253],"about_ca_topic_score_codex":0.0049352613,"about_ca_topic_score_gemma":0.00672954,"teacher_disagreement_score":0.0049352613,"about_ca_system_score_codex":0.0012959982,"about_ca_system_score_gemma":0.0014667846,"threshold_uncertainty_score":0.010474503},"labels":[],"label_agreement":null},{"id":"W4387143047","doi":"10.1109/re57278.2023.00020","title":"A Data-Driven Approach for Finding Requirements Relevant Feedback from TikTok and YouTube","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Metadata; Computer science; Social media; Focus (optics); Multimedia; World Wide Web; User requirements document; User engagement; Feature (linguistics); Human–computer interaction; Software engineering","score_opus":0.13927590013914581,"score_gpt":0.3365486849816968,"score_spread":0.197272784842551,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387143047","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.337006,0.0022180767,0.607557,0.0023391661,0.00025692908,0.002672204,0.03020843,0.010588516,0.0071537127],"genre_scores_gemma":[0.6654624,0.00042359543,0.2918506,0.0003315471,0.000075127595,0.0015543014,0.035817415,0.00023940494,0.004245596],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99771774,0.0005510711,0.00020375401,0.0005198957,0.00083558506,0.00017202258],"domain_scores_gemma":[0.9929249,0.003443663,0.0005989302,0.00037208482,0.0024409227,0.00021948975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001962203,0.0015144395,0.00068779645,0.004970287,0.0004842491,0.0009308604,0.0013244364,0.0010453924,0.0013676358],"category_scores_gemma":[0.01185049,0.00037113222,0.0009150778,0.002365954,0.00043796666,0.0013374055,0.0011226673,0.0011390755,0.0009200497],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001774292,0.0014441996,0.06352529,0.002259399,0.0003169251,0.0024053263,0.0023089414,0.058715425,0.062168237,0.0047124824,0.030066503,0.770303],"study_design_scores_gemma":[0.00008092004,0.000485367,0.019180229,0.00014640267,0.00009955063,0.00033253714,0.0013109792,0.93905085,0.0200896,0.005410335,0.013736193,0.00007698107],"about_ca_topic_score_codex":0.022044288,"about_ca_topic_score_gemma":0.048214108,"teacher_disagreement_score":0.022044288,"about_ca_system_score_codex":0.0018720663,"about_ca_system_score_gemma":0.0015091412,"threshold_uncertainty_score":0.043832004},"labels":[],"label_agreement":null},{"id":"W4387143077","doi":"10.1109/rew57809.2023.00021","title":"Mining Reddit Data to Elicit Students' Requirements During COVID-19 Pandemic","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Software requirements; Requirements elicitation; Requirements engineering; Event (particle physics); Identification (biology); Software; Requirement; Product (mathematics); Software engineering; Benchmarking; World Wide Web; Software development; Data science; Component-based software engineering","score_opus":0.2124887956126337,"score_gpt":0.43033092485876945,"score_spread":0.21784212924613575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387143077","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93329525,0.00032406976,0.04347725,0.0019346385,0.00008569054,0.0010924428,0.012618417,0.0012002441,0.0059719607],"genre_scores_gemma":[0.89817494,0.00027231406,0.074692704,0.0004549092,0.00004518297,0.001620345,0.02163233,0.00018575203,0.0029214928],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9892874,0.005534951,0.001038226,0.001086807,0.002398726,0.00065382186],"domain_scores_gemma":[0.91198546,0.062449224,0.006368642,0.0049128714,0.0124656465,0.0018181978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006879118,0.0008946698,0.0006567523,0.0038330078,0.00096384116,0.0013527548,0.0012230648,0.0017389372,0.0014081609],"category_scores_gemma":[0.051051084,0.00045492544,0.0006009008,0.002218297,0.0005831584,0.0015204557,0.0020722225,0.0018762964,0.0011099138],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016275677,0.0031889451,0.38524887,0.004809537,0.0002450821,0.0069818664,0.04047375,0.05064408,0.052324843,0.0040312735,0.038460426,0.4119638],"study_design_scores_gemma":[0.00024283308,0.0026500507,0.41967523,0.0014654893,0.00020056515,0.0019623134,0.06742879,0.30460468,0.06754469,0.0076185064,0.12601793,0.0005888426],"about_ca_topic_score_codex":0.007659696,"about_ca_topic_score_gemma":0.014419858,"teacher_disagreement_score":0.007659696,"about_ca_system_score_codex":0.0016808866,"about_ca_system_score_gemma":0.002032863,"threshold_uncertainty_score":0.03638071},"labels":[],"label_agreement":null},{"id":"W4387143087","doi":"10.1109/re57278.2023.00011","title":"User Driven Functionality Deletion for Mobile Apps","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; York University","funders":"","keywords":"Computer science; Usability; Mobile apps; World Wide Web; Software; Domain (mathematical analysis); User interface; Resource (disambiguation); Resource consumption; App store; Maintainability; Mobile device; Human–computer interaction; Software engineering; Operating system","score_opus":0.02892104170237283,"score_gpt":0.2952857328740373,"score_spread":0.2663646911716645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387143087","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9810581,0.00095791416,0.013085919,0.00038258912,0.00004047045,0.00030847828,0.0007174257,0.0015437012,0.0019054549],"genre_scores_gemma":[0.9784703,0.00022723897,0.018962642,0.00012728303,0.000032962154,0.000097765114,0.0009472142,0.000066601264,0.0010678164],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9957487,0.0015919886,0.0003696082,0.00055170414,0.0016177992,0.00012020347],"domain_scores_gemma":[0.9485383,0.036087226,0.006145619,0.0019866726,0.0066963555,0.00054583116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035733432,0.0008104348,0.0006171233,0.002313774,0.00056409725,0.0012015644,0.00067554566,0.00068926887,0.00070170517],"category_scores_gemma":[0.034991562,0.00034750908,0.00055448123,0.0010313586,0.00031661254,0.0012972591,0.00063674804,0.0006881715,0.0005399925],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015243018,0.00072359666,0.5304085,0.001721434,0.00032422904,0.0011487412,0.006651026,0.007084523,0.035149843,0.00040490407,0.0078127105,0.40704623],"study_design_scores_gemma":[0.00007742909,0.0032605876,0.8027838,0.00028461422,0.00039285413,0.0028945005,0.0028710691,0.14355305,0.026120367,0.0010591365,0.016452137,0.0002505372],"about_ca_topic_score_codex":0.0030963966,"about_ca_topic_score_gemma":0.008804982,"teacher_disagreement_score":0.0035733432,"about_ca_system_score_codex":0.000639497,"about_ca_system_score_gemma":0.0005944893,"threshold_uncertainty_score":0.018897891},"labels":[],"label_agreement":null},{"id":"W4387430331","doi":"10.1007/978-3-031-45275-8_36","title":"Unsupervised Graph Neural Networks for Source Code Similarity Detection","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; IBM (Canada); Polytechnique Montréal","funders":"","keywords":"Computer science; Source code; Inference; Similarity (geometry); Artificial intelligence; Graph; Unsupervised learning; Encoder; Artificial neural network; Code (set theory); Pattern recognition (psychology); Context (archaeology); Data mining; Theoretical computer science; Programming language","score_opus":0.027613856604873862,"score_gpt":0.2597452961992475,"score_spread":0.23213143959437366,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387430331","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04038147,0.0018134747,0.94737077,0.0002751992,0.00013095938,0.00008394791,0.0007744154,0.005522432,0.0036472962],"genre_scores_gemma":[0.4454909,0.0011437968,0.5327174,0.00019939773,0.00016812715,0.00015901709,0.0032466429,0.0007222469,0.01615247],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996295,0.00007181372,0.00001814489,0.00012904583,0.000114814604,0.000036788617],"domain_scores_gemma":[0.9990777,0.00043351782,0.00009804583,0.00014776367,0.00021443279,0.00002867361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00037079392,0.00070483604,0.00064813695,0.0019435142,0.00035377906,0.0008090945,0.0014576967,0.0011040196,0.0026586233],"category_scores_gemma":[0.001989046,0.00039996637,0.000646719,0.0019476839,0.00033890494,0.0012553643,0.00083332683,0.0010946409,0.0014916088],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016685418,0.00015159321,0.0016682651,0.0001407439,0.00009527968,0.000071843555,0.000043638658,0.13516144,0.016522387,0.0057996577,0.009300351,0.83087796],"study_design_scores_gemma":[0.00000386132,0.000018366805,0.00056708173,0.000009451366,0.000012452944,0.00003276549,0.000011997289,0.989286,0.0034188398,0.0056450814,0.0009882919,0.000005722192],"about_ca_topic_score_codex":0.006426096,"about_ca_topic_score_gemma":0.01346526,"teacher_disagreement_score":0.006426096,"about_ca_system_score_codex":0.00068757817,"about_ca_system_score_gemma":0.0005260665,"threshold_uncertainty_score":0.012777388},"labels":[],"label_agreement":null},{"id":"W4387606803","doi":"10.1145/3584931.3608438","title":"LLMs and the Infrastructure of CSCW","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Computer-supported cooperative work; Computer science; World Wide Web; Software engineering; Knowledge management; Human–computer interaction; Data science; Engineering; Work (physics)","score_opus":0.007774188708816054,"score_gpt":0.24041087359417962,"score_spread":0.23263668488536357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387606803","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022929119,0.0052585043,0.3343295,0.30767262,0.0011085531,0.00058340875,0.00047238357,0.0047131917,0.32293266],"genre_scores_gemma":[0.65446573,0.008570852,0.23801002,0.015711954,0.0012196044,0.0028334644,0.00130666,0.0015750059,0.07630676],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96789604,0.021646991,0.0018423698,0.0019088669,0.005419642,0.0012860631],"domain_scores_gemma":[0.9209093,0.042459518,0.0034594054,0.020426285,0.009052731,0.0036927268],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02553221,0.00073011575,0.0007123676,0.005816639,0.0065770517,0.019880528,0.003416805,0.0054588863,0.015242938],"category_scores_gemma":[0.07257922,0.0010768089,0.0010603868,0.0046322984,0.018777018,0.03449959,0.018273331,0.0063755605,0.0027411024],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00000955426,0.000029868928,0.00046213611,0.00010306761,0.000009091344,0.000035587942,0.0044837333,0.001445131,0.000110404486,0.9417218,0.01161197,0.039977737],"study_design_scores_gemma":[0.000028230155,0.000030850737,0.00045064685,0.00055321085,0.000010503835,0.00003868819,0.0033047344,0.0055829175,0.00036294575,0.70676136,0.28284058,0.00003536913],"about_ca_topic_score_codex":0.022587577,"about_ca_topic_score_gemma":0.01430883,"teacher_disagreement_score":0.029388413,"about_ca_system_score_codex":0.029388413,"about_ca_system_score_gemma":0.02364212,"threshold_uncertainty_score":0.21322888},"labels":[],"label_agreement":null},{"id":"W4387700648","doi":"10.1145/3617946.3617952","title":"An Interview with Gail Murphy - 2023 SIGSOFT Awardee","year":2023,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Vice president; Work (physics); Management; Productivity; Engineering management; Computer Science and Engineering; Engineering ethics; Library science; Engineering; Software; Software engineering; Medical education; Computer science; Sociology; Medicine","score_opus":0.02860339735246986,"score_gpt":0.27062594158016606,"score_spread":0.2420225442276962,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387700648","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027628276,0.0037399381,0.000741056,0.9234628,0.011138815,0.00008555323,0.00015308664,0.00015767418,0.032892805],"genre_scores_gemma":[0.25103575,0.0044582696,0.0018196999,0.5334571,0.0033561727,0.00035589025,0.00023015658,0.00034846427,0.20493835],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.993814,0.0029057667,0.00020195858,0.0004261797,0.0015539408,0.0010980653],"domain_scores_gemma":[0.98429525,0.004695391,0.0005839984,0.00031582985,0.003528436,0.0065811],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077545354,0.0005899492,0.0006000555,0.0021471381,0.016136989,0.008502217,0.0013913376,0.007479845,0.011640074],"category_scores_gemma":[0.022033405,0.0006364893,0.00038138806,0.0016985954,0.003441772,0.0065084244,0.0036881645,0.012983397,0.0045257416],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002479654,0.00009465357,0.0020305286,0.000029720173,0.000004613902,0.00048084822,0.020920044,0.00004755261,0.00025516402,0.003452343,0.95961696,0.013042726],"study_design_scores_gemma":[0.000010270434,0.00008167844,0.003087171,0.00012694523,0.0000046710356,0.0005552587,0.07606711,0.00020037181,0.00014307137,0.0012012029,0.9184617,0.00006053957],"about_ca_topic_score_codex":0.06555564,"about_ca_topic_score_gemma":0.09563074,"teacher_disagreement_score":0.06555564,"about_ca_system_score_codex":0.010558925,"about_ca_system_score_gemma":0.008778022,"threshold_uncertainty_score":0.13034809},"labels":[],"label_agreement":null},{"id":"W4387723770","doi":"10.48550/arxiv.2310.09575","title":"Common Challenges of Deep Reinforcement Learning Applications Development: An Empirical Study","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Consortium de Recherche et d’innovation en Aérospatiale au Québec; Canadian Institute for Advanced Research","keywords":"Computer science; Reinforcement learning; Taxonomy (biology); Popularity; Leverage (statistics); Artificial intelligence; Data science; Software engineering; Political science","score_opus":0.15514129597793713,"score_gpt":0.2693460052794603,"score_spread":0.11420470930152316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387723770","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9973411,0.0001998508,0.0010824681,0.00022819173,0.000006451238,0.000081254155,0.00014227757,0.000038032314,0.00088041596],"genre_scores_gemma":[0.99630654,0.00027115908,0.0022678606,0.00012086734,0.000012871361,0.0001108031,0.0003513037,0.000034536595,0.00052415626],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9903866,0.003415463,0.0011666315,0.0011555079,0.003237517,0.00063838105],"domain_scores_gemma":[0.8217056,0.12047485,0.02912013,0.0059384517,0.019150008,0.0036110089],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009320072,0.00041531358,0.0004156025,0.0026534153,0.0012364892,0.0018093521,0.0009167106,0.0010644337,0.0009932938],"category_scores_gemma":[0.08354793,0.00046704127,0.00035710644,0.0026562547,0.0012046793,0.0035895342,0.002015277,0.0016062289,0.00047956733],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037762674,0.0013586092,0.8633225,0.00079670956,0.000071997885,0.0010710164,0.028998582,0.0013336802,0.00270143,0.0009391633,0.0034934958,0.09553514],"study_design_scores_gemma":[0.000055237102,0.0012672635,0.90130967,0.00052837527,0.00007943036,0.0019223454,0.05156536,0.024131741,0.0034703836,0.0012802259,0.014260326,0.00012959784],"about_ca_topic_score_codex":0.003687215,"about_ca_topic_score_gemma":0.006585799,"teacher_disagreement_score":0.009320072,"about_ca_system_score_codex":0.0012649875,"about_ca_system_score_gemma":0.0017107886,"threshold_uncertainty_score":0.049289823},"labels":[],"label_agreement":null},{"id":"W4387739952","doi":"10.1007/s10664-023-10382-z","title":"Studying the characteristics of AIOps projects on GitHub","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Baseline (sea); Sample (material); Context (archaeology); Data science; Set (abstract data type); Quality (philosophy); Software engineering; Software; Open source; Anomaly detection; Data mining","score_opus":0.05649726322564533,"score_gpt":0.29934722376021994,"score_spread":0.2428499605345746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387739952","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9965515,0.00004424418,0.00018103595,0.00008670367,0.0000032528624,0.000021583692,0.00015026152,0.000030631647,0.0029308072],"genre_scores_gemma":[0.99712783,0.00007409049,0.00041810574,0.00002496528,0.0000066614566,0.000031844924,0.0004992452,0.000051341045,0.0017659497],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9971257,0.0007610168,0.00017625783,0.0003196046,0.001052457,0.00056500756],"domain_scores_gemma":[0.9633774,0.013354419,0.009690589,0.0016434967,0.0059043425,0.0060296613],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0023103412,0.00028106396,0.00020199611,0.0048844414,0.0010439756,0.0026170711,0.0007704103,0.00045689804,0.0032711786],"category_scores_gemma":[0.027060652,0.00022999039,0.00017243077,0.0073055234,0.0007731764,0.0020954462,0.001983269,0.00084047124,0.00086264435],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020136836,0.00036891858,0.9400247,0.000097630626,0.000044611214,0.0004916952,0.00970639,0.0008172618,0.002291508,0.0014023887,0.0022119556,0.042341597],"study_design_scores_gemma":[0.0000074770915,0.00013666271,0.9820157,0.000027864822,0.000009028285,0.00019085275,0.011849159,0.0018830816,0.00049666816,0.0002876472,0.0030780565,0.000017895672],"about_ca_topic_score_codex":0.016545564,"about_ca_topic_score_gemma":0.0351885,"teacher_disagreement_score":0.9951156,"about_ca_system_score_codex":0.0017122837,"about_ca_system_score_gemma":0.002061564,"threshold_uncertainty_score":0.032898545},"labels":[],"label_agreement":null},{"id":"W4387869001","doi":"10.1145/3630009","title":"The Good, the Bad, and the Missing: Neural Code Generation for Machine Learning Tasks","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Snippet; Artificial intelligence; Code (set theory); Code generation; Machine learning; Artificial neural network; Construct (python library); Set (abstract data type); Programming language; Natural language processing","score_opus":0.10990026194072332,"score_gpt":0.33788673840129657,"score_spread":0.22798647646057324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387869001","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53225744,0.009701511,0.4262698,0.0042275484,0.00043004417,0.0004561643,0.0015269882,0.015477096,0.009653416],"genre_scores_gemma":[0.7354818,0.0015009484,0.25550506,0.0007565064,0.00008323233,0.0004297448,0.0028238953,0.0006267389,0.0027921475],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9968021,0.0014979339,0.00024360412,0.00056551356,0.00072766305,0.00016320779],"domain_scores_gemma":[0.9866261,0.009543783,0.00076598045,0.0014828017,0.0013181923,0.0002630992],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041296105,0.0010442173,0.0005683819,0.0013173411,0.0005649148,0.0012346414,0.0017305302,0.0014082105,0.0011788377],"category_scores_gemma":[0.021779567,0.00038056145,0.00065346947,0.0011591071,0.0011058542,0.0028582963,0.0013751588,0.0021661092,0.00060907705],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096134737,0.0005498669,0.016872736,0.001091159,0.0001668452,0.000297874,0.0004746491,0.26278526,0.010054864,0.009590357,0.015297879,0.6818571],"study_design_scores_gemma":[0.00010104086,0.0002597223,0.0026518821,0.00009743898,0.000059071437,0.00012610875,0.00009445785,0.969383,0.009819735,0.013327458,0.0040429775,0.00003709741],"about_ca_topic_score_codex":0.0056711156,"about_ca_topic_score_gemma":0.0080198385,"teacher_disagreement_score":0.0056711156,"about_ca_system_score_codex":0.00155198,"about_ca_system_score_gemma":0.0015435441,"threshold_uncertainty_score":0.021839678},"labels":[],"label_agreement":null},{"id":"W4387914156","doi":"10.1109/codit58514.2023.10284124","title":"The Impact of Grid Search on Bug Resolution Prediction for Open-Source Software","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Eclipse; Computer science; Context (archaeology); Software bug; Decision tree; Support vector machine; Machine learning; Open source; Software; Data mining; Grid; Set (abstract data type); Random forest; Artificial intelligence; Open source software; Programming language","score_opus":0.04710568569820274,"score_gpt":0.3426597855985492,"score_spread":0.2955540999003465,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387914156","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97011876,0.00554824,0.018470485,0.00059836596,0.000107496475,0.000052988096,0.000925544,0.0025766036,0.001601492],"genre_scores_gemma":[0.98372114,0.00045886196,0.01417021,0.0000705473,0.000028133161,0.000020379994,0.0012549933,0.000053052383,0.00022264599],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970855,0.0010535222,0.00032295185,0.0006704963,0.00070726627,0.00016030497],"domain_scores_gemma":[0.96804976,0.02389683,0.0029373115,0.002540478,0.0021503381,0.0004253025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045089107,0.00058385386,0.0009762862,0.0022261664,0.00044838677,0.001067409,0.00083801214,0.00067958405,0.00045710042],"category_scores_gemma":[0.03463425,0.00018427098,0.00046996752,0.002034692,0.00039353225,0.0022487328,0.0008100955,0.00095247576,0.00035266145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013885075,0.00068946654,0.39263546,0.00067897193,0.00032402686,0.00026994955,0.00034383053,0.13499276,0.0034240398,0.0009139014,0.008587115,0.45575202],"study_design_scores_gemma":[0.000104991566,0.0010029699,0.094959155,0.0001562603,0.00012095535,0.0004437577,0.00047861997,0.8923983,0.0050942106,0.002663562,0.002528466,0.000048844337],"about_ca_topic_score_codex":0.005613903,"about_ca_topic_score_gemma":0.0063972617,"teacher_disagreement_score":0.005613903,"about_ca_system_score_codex":0.000513355,"about_ca_system_score_gemma":0.0008157363,"threshold_uncertainty_score":0.023845673},"labels":[],"label_agreement":null},{"id":"W4388185423","doi":"10.48175/ijarsct-13185","title":"Machine Learning-Powered Identification of Source Code Vulnerabilities","year":2023,"lang":"en","type":"article","venue":"International Journal of Advanced Research in Science Communication and Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Source code; Secure coding; Code review; Identification (biology); False positive paradox; Static program analysis; Scalability; Software; Code (set theory); Open source; Software quality; Software engineering; Software security assurance; Computer security; Software development; Database; Machine learning; Operating system; Programming language; Information security","score_opus":0.04525898680387183,"score_gpt":0.40541519087346245,"score_spread":0.3601562040695906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388185423","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05928114,0.0030832237,0.91424906,0.0009936043,0.00030238178,0.00036208177,0.001374048,0.011999327,0.008355127],"genre_scores_gemma":[0.5442096,0.0017583783,0.4422449,0.00083892775,0.0003841554,0.0004208456,0.0026223008,0.0007612191,0.0067597525],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964887,0.0007406471,0.00020383652,0.0008800419,0.0014977758,0.00018893661],"domain_scores_gemma":[0.980186,0.010192145,0.0035214007,0.0031302592,0.0027062225,0.00026399628],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002562177,0.0012339654,0.0012032223,0.0067143687,0.0005837142,0.0015153815,0.0021448631,0.0014678815,0.0028888518],"category_scores_gemma":[0.017500585,0.00039856153,0.0009339637,0.0027710185,0.0011931299,0.00304571,0.001689052,0.001891702,0.002872274],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014581632,0.00037348593,0.027303886,0.0007828903,0.00021772966,0.00013615367,0.0002135467,0.027602961,0.015273698,0.0073281126,0.008493666,0.91212815],"study_design_scores_gemma":[0.000045774555,0.00024596514,0.02158245,0.00030443733,0.00012980044,0.0007790489,0.0001451406,0.89028925,0.035864893,0.03488554,0.015615848,0.00011195304],"about_ca_topic_score_codex":0.0010778178,"about_ca_topic_score_gemma":0.0018168949,"teacher_disagreement_score":0.0067143687,"about_ca_system_score_codex":0.0007479004,"about_ca_system_score_gemma":0.0011827084,"threshold_uncertainty_score":0.0135502815},"labels":[],"label_agreement":null},{"id":"W4388190842","doi":"10.2139/ssrn.4615254","title":"A First Look at Information Highlighting in Stack Overflow Answers","year":2023,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of Manitoba","funders":"","keywords":"Stack (abstract data type); Computer science; Data science; Programming language","score_opus":0.013266734943572705,"score_gpt":0.24752863146605247,"score_spread":0.23426189652247978,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388190842","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31832227,0.01953841,0.33877715,0.031422548,0.0025991607,0.0005455314,0.0067478907,0.020007713,0.26203933],"genre_scores_gemma":[0.82910514,0.0042751343,0.09886676,0.0039366283,0.0015289059,0.00009896559,0.0029694804,0.0033745572,0.055844426],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9960522,0.0009631165,0.000244167,0.00037999908,0.0017405525,0.0006200245],"domain_scores_gemma":[0.9766822,0.01555629,0.001379057,0.0017829464,0.0041104592,0.00048909377],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029058133,0.0006361539,0.0008160772,0.0051840097,0.0025810783,0.0054863286,0.0013197063,0.0028211274,0.031614244],"category_scores_gemma":[0.034223326,0.0005380274,0.0006646247,0.0052598617,0.0020539528,0.01414801,0.0031117406,0.0025212576,0.0047002323],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020824128,0.0003368767,0.022426616,0.002213789,0.000107224565,0.004635714,0.021660682,0.004330699,0.04367969,0.28926432,0.107940696,0.5013213],"study_design_scores_gemma":[0.00011707974,0.00070068263,0.018692143,0.0017847763,0.00024836915,0.0057318113,0.014729648,0.032693464,0.07050043,0.22631155,0.62811404,0.00037601928],"about_ca_topic_score_codex":0.0037255173,"about_ca_topic_score_gemma":0.00376043,"teacher_disagreement_score":0.031614244,"about_ca_system_score_codex":0.0013608374,"about_ca_system_score_gemma":0.001331107,"threshold_uncertainty_score":0.10576016},"labels":[],"label_agreement":null},{"id":"W4388474969","doi":"10.18280/ria.370515","title":"Using Natural Language Processing for Programming Language Code Classification with Multinomial Naive Bayes","year":2023,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Natural language processing; Multinomial distribution; Naive Bayes classifier; Artificial intelligence; Code (set theory); Programming language; Statistics; Mathematics; Support vector machine","score_opus":0.07372243200853963,"score_gpt":0.3505859901875498,"score_spread":0.2768635581790102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388474969","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037645537,0.00092648336,0.9475987,0.0009197662,0.00024876863,0.0006683058,0.0016772278,0.0072641424,0.0030509771],"genre_scores_gemma":[0.19904618,0.00037301934,0.7908959,0.0005453856,0.0001991401,0.0009710726,0.0054190704,0.00034040178,0.0022097947],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9948757,0.0019927667,0.00057727733,0.001254541,0.0010823156,0.00021745503],"domain_scores_gemma":[0.98933613,0.007538313,0.0006985217,0.00064121286,0.0016117862,0.00017404568],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055639013,0.0016285495,0.0011683728,0.005710367,0.001317609,0.002584539,0.0020491339,0.0015917829,0.0035258755],"category_scores_gemma":[0.021514079,0.0006331345,0.0016946726,0.002905337,0.0008726656,0.0024403993,0.0012029181,0.0022452723,0.0027193052],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057018606,0.00059827685,0.021307958,0.00078077323,0.00021433244,0.0004475425,0.00079818995,0.06390419,0.006978954,0.0107906815,0.019903755,0.8737053],"study_design_scores_gemma":[0.000078039804,0.00010022708,0.0022819715,0.00015752797,0.000059145365,0.00029125312,0.00027683718,0.9546323,0.0039356314,0.030801242,0.0073342,0.0000516098],"about_ca_topic_score_codex":0.009927927,"about_ca_topic_score_gemma":0.012095787,"teacher_disagreement_score":0.009927927,"about_ca_system_score_codex":0.0015199895,"about_ca_system_score_gemma":0.0026826984,"threshold_uncertainty_score":0.029425085},"labels":[],"label_agreement":null},{"id":"W4388483163","doi":"10.1109/secdev56634.2023.00016","title":"Grading on a Curve: How Rust can Facilitate New Contributors while Decreasing Vulnerabilities","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Canada Research Chairs","keywords":"Rust (programming language); Implementation; Computer science; Grading (engineering); Code (set theory); Computer security; Data science; World Wide Web; Risk analysis (engineering); Software engineering; Engineering; Business; Programming language","score_opus":0.059685232835197784,"score_gpt":0.27329409456351145,"score_spread":0.21360886172831367,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388483163","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82382506,0.0020021219,0.13620263,0.0028835337,0.00030591735,0.00093878555,0.0044167927,0.0054723253,0.023952784],"genre_scores_gemma":[0.9474989,0.0002931702,0.04516428,0.00022366794,0.00008154749,0.00034532492,0.0021250986,0.00046893788,0.0037989693],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9847474,0.0056611225,0.0010060191,0.0030049197,0.0044564623,0.0011240382],"domain_scores_gemma":[0.80899787,0.12638941,0.020810027,0.0174467,0.020374494,0.0059815333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02223005,0.0010685099,0.001191369,0.0072572515,0.001653399,0.0044522057,0.0022133559,0.0022611464,0.0061799907],"category_scores_gemma":[0.1904505,0.0004936337,0.001160071,0.0044946484,0.0022741973,0.007831783,0.0029476434,0.0025452892,0.0034880952],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013516198,0.0006904867,0.47656858,0.00072561885,0.00031895368,0.0004112358,0.0044483575,0.050147798,0.0027979517,0.011562947,0.02395403,0.42702252],"study_design_scores_gemma":[0.00030569267,0.0027913582,0.41364816,0.0007462108,0.00033385234,0.0016132579,0.0065080514,0.45794758,0.010631789,0.051595863,0.05331802,0.0005600976],"about_ca_topic_score_codex":0.0076647648,"about_ca_topic_score_gemma":0.007974649,"teacher_disagreement_score":0.02223005,"about_ca_system_score_codex":0.0021620276,"about_ca_system_score_gemma":0.0018280343,"threshold_uncertainty_score":0.117565095},"labels":[],"label_agreement":null},{"id":"W4388483496","doi":"10.1109/ase56229.2023.00178","title":"iASTMapper: An Iterative Similarity-Based Abstract Syntax Tree Mapping Algorithm","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Source code; Code (set theory); Syntax; Algorithm; Tree (set theory); Node (physics); Similarity (geometry); Heuristic; Theoretical computer science; Programming language; Artificial intelligence; Mathematics; Set (abstract data type); Image (mathematics)","score_opus":0.04030030330896766,"score_gpt":0.29348524823926203,"score_spread":0.25318494493029436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388483496","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0141066015,0.00023972841,0.96869856,0.0001268292,0.00008546679,0.0002460865,0.0002207776,0.01480208,0.0014739351],"genre_scores_gemma":[0.059626665,0.00009721257,0.9354462,0.00014775104,0.000028063887,0.00024196404,0.0009729382,0.0009791201,0.002460014],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982699,0.00026664074,0.00016273577,0.00047201422,0.0006696824,0.00015910841],"domain_scores_gemma":[0.99760497,0.00075515633,0.00025135186,0.0005170619,0.0007703358,0.000100955964],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013837994,0.0018025803,0.0014863089,0.00276758,0.0010232993,0.0014492472,0.0034418148,0.0017543582,0.0044543515],"category_scores_gemma":[0.006862942,0.0009046243,0.0018912337,0.0021187873,0.0009853129,0.0029639192,0.002853788,0.0020256417,0.0029002123],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027738814,0.0001928255,0.0035349978,0.00026830033,0.00012775142,0.00026669877,0.0005962946,0.06143771,0.021504857,0.008756157,0.014326093,0.8887108],"study_design_scores_gemma":[0.00013180445,0.00018573177,0.0010985925,0.000035738634,0.000072997675,0.00039099436,0.00023686007,0.9483268,0.018836819,0.017232785,0.013386169,0.00006471357],"about_ca_topic_score_codex":0.005638758,"about_ca_topic_score_gemma":0.007469112,"teacher_disagreement_score":0.005638758,"about_ca_system_score_codex":0.0008194119,"about_ca_system_score_gemma":0.0026970329,"threshold_uncertainty_score":0.01490128},"labels":[],"label_agreement":null},{"id":"W4388502412","doi":"10.1109/ase56229.2023.00030","title":"Repeated Builds During Code Review: An Empirical Study of the OpenStack Community","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Australian Research Council","keywords":"Computer science; Software deployment; Code (set theory); Set (abstract data type); Process (computing); Empirical research; Software engineering; Programming language; Computer security","score_opus":0.09380204085335725,"score_gpt":0.38939035847685866,"score_spread":0.2955883176235014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388502412","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9976774,0.00032851915,0.000543906,0.0002522399,0.000013526619,0.0000759938,0.00015134344,0.0000622298,0.0008948265],"genre_scores_gemma":[0.99700636,0.00022169133,0.0012511086,0.00020032155,0.00003800828,0.000112714304,0.0004342394,0.00009207522,0.0006435022],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9614832,0.019014722,0.0023914017,0.0044116518,0.011095409,0.0016036229],"domain_scores_gemma":[0.4571471,0.37265784,0.094706714,0.018951558,0.04678658,0.009750153],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.02938501,0.00061605143,0.00067582383,0.007145508,0.0027054893,0.002726057,0.0024017617,0.0019793247,0.0014899316],"category_scores_gemma":[0.26200053,0.0007281233,0.00048144106,0.0046119224,0.0030334122,0.005282151,0.0028919945,0.0024684938,0.0008811384],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008199166,0.0014889436,0.81300247,0.0009300601,0.00030066905,0.0023209688,0.08839908,0.0014913566,0.0037554845,0.00087402825,0.008044184,0.07857289],"study_design_scores_gemma":[0.00013294617,0.001478307,0.9071565,0.00048669608,0.00011328052,0.0029201154,0.055846702,0.012608944,0.0023688476,0.0010181866,0.015581228,0.00028812958],"about_ca_topic_score_codex":0.007243331,"about_ca_topic_score_gemma":0.011932893,"teacher_disagreement_score":0.97061497,"about_ca_system_score_codex":0.0017705053,"about_ca_system_score_gemma":0.0023434802,"threshold_uncertainty_score":0.15540469},"labels":[],"label_agreement":null},{"id":"W4388758252","doi":"10.1109/tse.2023.3332568","title":"Properties and Styles of Software Technology Tutorials","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; JavaScript; Python (programming language); Documentation; Software; TypeScript; Java; Resource (disambiguation); World Wide Web; Software engineering; Information retrieval; Programming language","score_opus":0.025863080000366374,"score_gpt":0.23587883220421202,"score_spread":0.21001575220384563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388758252","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8974189,0.00093420746,0.06918813,0.0002526372,0.00003527754,0.00026182225,0.003456474,0.0019727591,0.026479706],"genre_scores_gemma":[0.9694569,0.00025626103,0.023564598,0.000047350524,0.000036187692,0.0001960281,0.0032095013,0.00047472882,0.0027584508],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.98853475,0.0041256538,0.0016049455,0.0015334359,0.0035137169,0.0006875992],"domain_scores_gemma":[0.87407136,0.067504674,0.02351988,0.012456725,0.017671041,0.004776269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058185006,0.00046447155,0.0004356266,0.008416725,0.00095720816,0.004696299,0.00087617955,0.000719772,0.0031700053],"category_scores_gemma":[0.10379945,0.0005593051,0.00067100726,0.0066967797,0.0010552256,0.0054889387,0.0017192855,0.00079257536,0.0009182045],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009242467,0.0004120711,0.5665036,0.0011407346,0.00030914365,0.00084656634,0.014383565,0.01523705,0.021428967,0.08146887,0.0062105427,0.2911347],"study_design_scores_gemma":[0.0001271482,0.0009431941,0.6772908,0.00064046634,0.00033930043,0.0036040824,0.008518594,0.08952387,0.030250786,0.08725972,0.10110138,0.0004006696],"about_ca_topic_score_codex":0.0016057687,"about_ca_topic_score_gemma":0.0015466028,"teacher_disagreement_score":0.008416725,"about_ca_system_score_codex":0.0015556797,"about_ca_system_score_gemma":0.0010635625,"threshold_uncertainty_score":0.030771554},"labels":[],"label_agreement":null},{"id":"W4388768226","doi":"10.1002/smr.2639","title":"A catalog of metrics at source code level for vulnerability prediction: A systematic mapping study","year":2023,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Vulnerability (computing); Code review; Software quality; Software security assurance; Software; Data mining; Software metric; Predictive modelling; Quality (philosophy); Machine learning; Data science; Software development; Computer security; Information security","score_opus":0.06678564583861035,"score_gpt":0.3157488891498784,"score_spread":0.24896324331126807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388768226","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31184083,0.6101332,0.048179615,0.0033020133,0.00036982712,0.006262303,0.009619957,0.0003605294,0.009931755],"genre_scores_gemma":[0.6880297,0.23052229,0.06675463,0.0009349677,0.00013160077,0.0063236034,0.0063114627,0.00014891019,0.000842811],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9769976,0.0071804626,0.0067775943,0.0020988039,0.006521446,0.00042404607],"domain_scores_gemma":[0.78314626,0.14221552,0.022698617,0.0070326277,0.04365261,0.0012544482],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.036409933,0.0010878096,0.0016862248,0.063259,0.0012954411,0.0028907037,0.0012544588,0.00091391196,0.0017299477],"category_scores_gemma":[0.13853504,0.0007714306,0.0024703213,0.039606996,0.000994512,0.0057410398,0.0029308046,0.0010935882,0.0003754368],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024950376,0.00026753987,0.14401598,0.14879838,0.002409174,0.000735378,0.010308613,0.0012569483,0.0027046637,0.006397526,0.006736355,0.6761199],"study_design_scores_gemma":[0.00018814333,0.0016284002,0.2558088,0.5196512,0.01759657,0.0026198905,0.029746117,0.0064011416,0.008636277,0.010413477,0.14697209,0.00033800522],"about_ca_topic_score_codex":0.004918696,"about_ca_topic_score_gemma":0.00987088,"teacher_disagreement_score":0.9635901,"about_ca_system_score_codex":0.0027692544,"about_ca_system_score_gemma":0.012575182,"threshold_uncertainty_score":0.19255644},"labels":[],"label_agreement":null},{"id":"W4388891009","doi":"10.48550/arxiv.2311.11177","title":"Assessing the Security of GitHub Copilot Generated Code -- A Targeted Replication Study","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Massey University; Canadian Institute for Advanced Research","keywords":"Computer science; Python (programming language); Code (set theory); Computer security; Replication (statistics); Software engineering; Programming language","score_opus":0.16602238256579852,"score_gpt":0.2735873181719801,"score_spread":0.1075649356061816,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388891009","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97092664,0.0006227695,0.015808035,0.00082392775,0.00011953469,0.0007623535,0.0011451078,0.006417732,0.0033737796],"genre_scores_gemma":[0.9600811,0.0003108678,0.029391114,0.00047194771,0.000047468107,0.00077316945,0.0038364253,0.002426361,0.002661548],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9854375,0.0047895457,0.00090126315,0.002103005,0.0062499684,0.0005187353],"domain_scores_gemma":[0.8499159,0.055384155,0.008945501,0.05046512,0.033655044,0.0016342701],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015125493,0.0011520326,0.00073807576,0.0033511831,0.0009617176,0.0016238071,0.0027379887,0.0017094266,0.0010112152],"category_scores_gemma":[0.10417705,0.00076642603,0.0013048688,0.0018599514,0.0022866388,0.0037299849,0.0026029735,0.0030613057,0.0009165613],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007179524,0.005157462,0.24862649,0.0037031076,0.001588329,0.0044598873,0.018866489,0.08569926,0.08826821,0.007746077,0.041159388,0.48754582],"study_design_scores_gemma":[0.0011827091,0.010009776,0.2431041,0.0009296993,0.0014709523,0.005437028,0.00537844,0.52146864,0.13327406,0.008300374,0.068772405,0.0006717807],"about_ca_topic_score_codex":0.0095826555,"about_ca_topic_score_gemma":0.0065908893,"teacher_disagreement_score":0.9848745,"about_ca_system_score_codex":0.0028992954,"about_ca_system_score_gemma":0.0019593104,"threshold_uncertainty_score":0.079992235},"labels":[],"label_agreement":null},{"id":"W4388948694","doi":"10.1016/j.jss.2023.111914","title":"Commit-time defect prediction using one-class classification","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Commit; Support vector machine; Computer science; Class (philosophy); Random forest; Binary number; Machine learning; Binary classification; Overhead (engineering); Software; Data mining; Artificial intelligence; Database; Mathematics","score_opus":0.05332738242996626,"score_gpt":0.2774134836601288,"score_spread":0.22408610123016254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388948694","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85644317,0.0010777517,0.13281271,0.00034567688,0.00049811136,0.00024280179,0.002531105,0.0035702437,0.0024783856],"genre_scores_gemma":[0.96357673,0.000163591,0.03068095,0.000044260796,0.00010246503,0.00009120593,0.0027808747,0.00009307188,0.002466776],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981982,0.00016243878,0.00021528099,0.0005157096,0.00063113746,0.0002773063],"domain_scores_gemma":[0.98492473,0.0065775607,0.0018273651,0.0015755802,0.00432246,0.00077228213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001759807,0.0013036764,0.0012928655,0.004183825,0.00067774043,0.0012387607,0.0017971017,0.0013475694,0.0017852494],"category_scores_gemma":[0.008322848,0.00022839906,0.00097385683,0.0017644361,0.00035748063,0.0014450724,0.00087933557,0.0013169516,0.0009809172],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025013655,0.0021799535,0.23877862,0.000367608,0.00032064778,0.0007049998,0.00022770381,0.047233928,0.011714808,0.0010285889,0.010780469,0.68416125],"study_design_scores_gemma":[0.000043075517,0.0004720723,0.029711325,0.000035185923,0.00012171933,0.0003530446,0.00013060347,0.95956826,0.0070676967,0.001316543,0.0011352438,0.000045240544],"about_ca_topic_score_codex":0.005491266,"about_ca_topic_score_gemma":0.005809347,"teacher_disagreement_score":0.005491266,"about_ca_system_score_codex":0.0004983546,"about_ca_system_score_gemma":0.0010433184,"threshold_uncertainty_score":0.010918617},"labels":[],"label_agreement":null},{"id":"W4388954848","doi":"10.1145/3605770.3625215","title":"Distinguishing AI- and Human-Generated Code: A Case Study","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Code (set theory); Task (project management); Code review; Software; Source code; Artificial intelligence; Programming language; Natural language processing; Software engineering; Computer security; Static program analysis; Software development; Engineering","score_opus":0.06023643670637761,"score_gpt":0.35497965548433363,"score_spread":0.294743218777956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388954848","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96685654,0.0004188417,0.026377197,0.00077101425,0.00006035368,0.0005029826,0.0006180994,0.000796802,0.0035980484],"genre_scores_gemma":[0.9491886,0.00021456035,0.047339723,0.0003299984,0.000030709376,0.00020313276,0.000978566,0.00015834667,0.0015563861],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9905409,0.0048566945,0.001122061,0.0012397077,0.0019322262,0.00030841588],"domain_scores_gemma":[0.8948419,0.08384417,0.0037021476,0.009290084,0.0071165236,0.0012052412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061747506,0.000622768,0.00039915982,0.0017606366,0.0013572236,0.0015107422,0.0012714923,0.002975075,0.0013055445],"category_scores_gemma":[0.0497611,0.00026564155,0.0005160863,0.0015760664,0.0021336349,0.0024324814,0.001833271,0.0015130226,0.00086746673],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024506461,0.007678123,0.2309566,0.0037331395,0.0003673668,0.033349466,0.07114043,0.025528373,0.07237092,0.0132130645,0.023919594,0.5152922],"study_design_scores_gemma":[0.0012385478,0.0062564206,0.2007444,0.0013802946,0.000368369,0.06046988,0.045330323,0.25871992,0.24812612,0.031828456,0.14496306,0.0005742897],"about_ca_topic_score_codex":0.004277987,"about_ca_topic_score_gemma":0.008619157,"teacher_disagreement_score":0.0061747506,"about_ca_system_score_codex":0.00090207916,"about_ca_system_score_gemma":0.00095407467,"threshold_uncertainty_score":0.032655656},"labels":[],"label_agreement":null},{"id":"W4389068312","doi":"10.1016/j.jss.2023.111907","title":"Software engineering practices for machine learning — Adoption, effects, and team assessment","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Software engineering; Software; Computer science; Engineering management; Knowledge management; Artificial intelligence; Engineering; Operating system","score_opus":0.0184189781801899,"score_gpt":0.29245880791695766,"score_spread":0.27403982973676777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389068312","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9411087,0.0148982685,0.02405472,0.0062334063,0.00014236629,0.00068373693,0.0002322122,0.00014584405,0.012500821],"genre_scores_gemma":[0.98528755,0.0032322386,0.010295781,0.00034872323,0.000035724785,0.00036645622,0.00008850744,0.00003218779,0.00031294994],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8701115,0.082075,0.012644788,0.005696317,0.027416766,0.0020555703],"domain_scores_gemma":[0.4722625,0.40090653,0.06349672,0.019364744,0.03944261,0.004526845],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09284354,0.0005338829,0.0006659377,0.0076722396,0.0017936733,0.0047389497,0.0014060035,0.0012968947,0.001162709],"category_scores_gemma":[0.2665867,0.00067156286,0.0012355555,0.006590874,0.003535646,0.0070881713,0.005555839,0.0017828583,0.0002803184],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002454828,0.0005159163,0.32622007,0.0077720247,0.0008586016,0.000520244,0.14185546,0.0024731131,0.0022685735,0.012473472,0.002437148,0.5023599],"study_design_scores_gemma":[0.00018184092,0.0028118382,0.70520747,0.022598648,0.0010078992,0.0011099983,0.17559801,0.008239176,0.0050240178,0.026328256,0.051506065,0.00038681162],"about_ca_topic_score_codex":0.0032576444,"about_ca_topic_score_gemma":0.0035850878,"teacher_disagreement_score":0.90715647,"about_ca_system_score_codex":0.005209069,"about_ca_system_score_gemma":0.005442846,"threshold_uncertainty_score":0.4910094},"labels":[],"label_agreement":null},{"id":"W4389141459","doi":"10.1007/s10664-023-10389-6","title":"Silent bugs in deep learning frameworks: an empirical study of Keras and TensorFlow","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Fonds de recherche du Québec; Consortium de Recherche et d’innovation en Aérospatiale au Québec; Canadian Institute for Advanced Research","keywords":"Software bug; Computer science; Relevance (law); Debugging; Artificial intelligence; Empirical research; Deep learning; Machine learning; Software; Programming language","score_opus":0.02542377729338102,"score_gpt":0.3172849195316602,"score_spread":0.2918611422382792,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389141459","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98428345,0.0008296527,0.011164901,0.0008243346,0.000039135208,0.00004250032,0.00018922972,0.0006543831,0.001972387],"genre_scores_gemma":[0.99583143,0.00008232419,0.003395723,0.00006705879,0.000012638521,0.000015785034,0.00017944306,0.000113328846,0.00030230856],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98871917,0.0046696966,0.00068742054,0.0014012884,0.0036139083,0.0009084697],"domain_scores_gemma":[0.68801606,0.2473794,0.024937013,0.023321927,0.012175107,0.004170527],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0170036,0.0006607831,0.0005869774,0.0017123846,0.0012889195,0.0018038333,0.002579601,0.0023034262,0.0020983776],"category_scores_gemma":[0.22011997,0.0006145296,0.00061695627,0.0019721857,0.0033000768,0.008320617,0.0021248509,0.00485606,0.00033032589],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0045130397,0.0067428173,0.53478426,0.0011737962,0.0006179828,0.0010561844,0.008186323,0.08544864,0.0048960587,0.08509899,0.019933395,0.2475485],"study_design_scores_gemma":[0.0005670531,0.002477652,0.121728145,0.00049396296,0.00043844318,0.0012563339,0.0042569195,0.7412241,0.0057173823,0.1147117,0.006912648,0.00021562219],"about_ca_topic_score_codex":0.007722823,"about_ca_topic_score_gemma":0.0083760675,"teacher_disagreement_score":0.9829964,"about_ca_system_score_codex":0.001782899,"about_ca_system_score_gemma":0.0025358687,"threshold_uncertainty_score":0.08992469},"labels":[],"label_agreement":null},{"id":"W4389141545","doi":"10.1007/s10664-023-10399-4","title":"Unreproducible builds: time to fix, causes, and correlation with external ecosystem factors","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software; Process (computing); Reproducibility; Cyclomatic complexity; Software engineering; Data science; Operating system; Statistics; Mathematics","score_opus":0.019738238128314236,"score_gpt":0.25718784017076846,"score_spread":0.23744960204245422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389141545","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9977914,0.00025564607,0.0007366951,0.00009732312,0.0000056098916,0.00001203276,0.00018873617,0.000022112477,0.00089042063],"genre_scores_gemma":[0.99938715,0.000054721895,0.0001992388,0.0000068903655,0.0000049596983,0.0000063006833,0.00010966199,0.000011952993,0.00021915077],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99698776,0.0009901237,0.00044777605,0.00043034676,0.000716276,0.00042762636],"domain_scores_gemma":[0.7642281,0.1638295,0.04883907,0.011619904,0.00571416,0.0057692127],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0065181684,0.00029339464,0.00034346743,0.002086558,0.00054855633,0.0019713982,0.0010098846,0.0011389848,0.00473625],"category_scores_gemma":[0.09834314,0.00048948114,0.0006231088,0.0016338846,0.0012772689,0.0021106217,0.0014358924,0.0021982014,0.0005216803],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017369473,0.00018653694,0.9921523,0.00002150272,0.00008233988,0.0001050491,0.00030061373,0.0016078675,0.00020199773,0.00048565824,0.00011022834,0.0045720963],"study_design_scores_gemma":[0.00001875145,0.0001607855,0.9894917,0.000030941548,0.000097782366,0.0004359326,0.00084907736,0.00553341,0.0005151288,0.0023240703,0.00051719666,0.000025243451],"about_ca_topic_score_codex":0.0053844987,"about_ca_topic_score_gemma":0.008559096,"teacher_disagreement_score":0.9934818,"about_ca_system_score_codex":0.00088695396,"about_ca_system_score_gemma":0.0016363972,"threshold_uncertainty_score":0.03447181},"labels":[],"label_agreement":null},{"id":"W4389158520","doi":"10.1145/3611643.3613087","title":"A Vision on Intentions in Software Engineering","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Toronto","funders":"","keywords":"Computer science; Sketch; Stakeholder; Software evolution; Software; Software development; Software engineering; Knowledge management; Data science; Human–computer interaction; Software construction","score_opus":0.019193505145080763,"score_gpt":0.2842701505134527,"score_spread":0.26507664536837194,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389158520","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012067328,0.012999386,0.79772127,0.08403711,0.0016229951,0.00020818065,0.00014267291,0.0006249558,0.090576164],"genre_scores_gemma":[0.48872656,0.014009915,0.46547422,0.013411842,0.0019508335,0.00084059493,0.00025834504,0.00044288105,0.014884832],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98159677,0.011169995,0.0011109404,0.002101093,0.0032248253,0.0007964232],"domain_scores_gemma":[0.9653544,0.02167067,0.002093006,0.0045898664,0.00429713,0.0019949006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020726485,0.0017183822,0.0009906763,0.005751105,0.0040715104,0.012513576,0.0032201232,0.009464322,0.0030756446],"category_scores_gemma":[0.021036316,0.0014267851,0.0024011198,0.003076326,0.041327022,0.032216914,0.0073708445,0.012992332,0.0014649167],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017696953,0.00002839564,0.00032408943,0.00012121487,0.00000918684,0.000039540315,0.0028046977,0.0010132949,0.00019991647,0.98311317,0.0013806397,0.01094819],"study_design_scores_gemma":[0.000017718512,0.000049150734,0.0002893425,0.00030257725,0.000017595368,0.000087016335,0.0012482353,0.0034695813,0.00025987026,0.95098734,0.043233234,0.000038216873],"about_ca_topic_score_codex":0.0051422403,"about_ca_topic_score_gemma":0.0023358993,"teacher_disagreement_score":0.020726485,"about_ca_system_score_codex":0.0060827434,"about_ca_system_score_gemma":0.005680922,"threshold_uncertainty_score":0.10961348},"labels":[],"label_agreement":null},{"id":"W4389159648","doi":"10.1145/3611643.3613082","title":"Towards Feature-Based Analysis of the Machine Learning Development Lifecycle","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Taxonomy (biology); Feature (linguistics); Software engineering; Trustworthiness; Feature engineering; Formal concept analysis; Software; Software development; Data science; Artificial intelligence; Machine learning; Deep learning; Programming language","score_opus":0.01892595550722531,"score_gpt":0.26554302143317676,"score_spread":0.24661706592595145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389159648","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020892322,0.00029770448,0.9753209,0.00053199765,0.000013060552,0.00014774999,0.00024213729,0.00087367464,0.0016805143],"genre_scores_gemma":[0.1772998,0.00034350864,0.82013416,0.000087726534,0.000026691194,0.00030326314,0.0008959497,0.00021278784,0.0006960802],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99543166,0.0014613849,0.00042444307,0.00057126547,0.0017328758,0.00037832814],"domain_scores_gemma":[0.97575563,0.010432912,0.0030145738,0.0046545337,0.00571399,0.00042853432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061631477,0.0009306681,0.0007738457,0.0070436434,0.00096260913,0.004611095,0.0018443107,0.0013305151,0.0012360452],"category_scores_gemma":[0.025310285,0.0007812395,0.002135014,0.0034432542,0.0016079434,0.006695635,0.002767926,0.0021462566,0.0005255721],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018019414,0.0004446114,0.048372425,0.00079261075,0.00018790543,0.0009670956,0.004969167,0.106719464,0.011078965,0.30124655,0.004916528,0.52012455],"study_design_scores_gemma":[0.000022816908,0.00013937599,0.008393493,0.00031565037,0.00008835473,0.0003291206,0.0011150152,0.6983347,0.009341822,0.2624786,0.01935617,0.000084936866],"about_ca_topic_score_codex":0.006022611,"about_ca_topic_score_gemma":0.0041179,"teacher_disagreement_score":0.0070436434,"about_ca_system_score_codex":0.0024651003,"about_ca_system_score_gemma":0.0032713525,"threshold_uncertainty_score":0.032594264},"labels":[],"label_agreement":null},{"id":"W4389208648","doi":"10.1145/3611643.3617851","title":"A Data Set of Extracted Rationale from Linux Kernel Commit Messages","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Commit; Computer science; Set (abstract data type); Kernel (algebra); Component (thermodynamics); Linux kernel; World Wide Web; Operating system; Programming language; Database","score_opus":0.11631883728867144,"score_gpt":0.33789378860420616,"score_spread":0.22157495131553473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389208648","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28544584,0.0016435775,0.031312045,0.0019199011,0.0005365936,0.0012785177,0.660008,0.0076097837,0.010245817],"genre_scores_gemma":[0.17347373,0.0004919736,0.06496229,0.00043764216,0.00009957175,0.0019964,0.7524721,0.0009761909,0.0050901063],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99680483,0.0007636111,0.000621133,0.0005652781,0.0011054586,0.00013972966],"domain_scores_gemma":[0.9496216,0.035990465,0.0035398148,0.0026571967,0.007495479,0.0006954934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00240582,0.0010941775,0.00043099176,0.0050341403,0.0013024086,0.0011600759,0.0008707254,0.0024424104,0.0035842448],"category_scores_gemma":[0.029182775,0.00049610686,0.0005335644,0.0032032474,0.0006896914,0.0014428508,0.0009920435,0.0017768617,0.0033356417],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003956593,0.0014725664,0.07402628,0.021072559,0.00033088695,0.009070124,0.019562066,0.008302806,0.12608315,0.011327806,0.41427138,0.3105238],"study_design_scores_gemma":[0.0006712637,0.0010458934,0.22069025,0.0018996244,0.0002916533,0.0035009882,0.007076093,0.025259105,0.075493015,0.009088245,0.6545013,0.00048250204],"about_ca_topic_score_codex":0.0075808926,"about_ca_topic_score_gemma":0.017068237,"teacher_disagreement_score":0.0075808926,"about_ca_system_score_codex":0.0011550835,"about_ca_system_score_gemma":0.002235199,"threshold_uncertainty_score":0.015073538},"labels":[],"label_agreement":null},{"id":"W4389260672","doi":"10.48550/arxiv.2311.18057","title":"Non Linear Software Documentation with Interactive Code Examples","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Computer science; Software documentation; Field (mathematics); Code (set theory); Navigability; World Wide Web; Software; Quality (philosophy); Information retrieval; Baseline (sea); Technical documentation; Multimedia; Software development; Software construction; Programming language; Set (abstract data type)","score_opus":0.08253020477204236,"score_gpt":0.2359919722441411,"score_spread":0.15346176747209872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389260672","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21894012,0.0016820998,0.7047921,0.0030433864,0.0002674671,0.0006609089,0.001164782,0.017463192,0.051986005],"genre_scores_gemma":[0.4666738,0.0006345092,0.5151909,0.00034787937,0.00009141469,0.0004873741,0.0014102894,0.0015126402,0.013651207],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99523526,0.0022467035,0.0005172588,0.00043215087,0.0013738849,0.000194711],"domain_scores_gemma":[0.9078298,0.06225517,0.005851654,0.015168132,0.0074971663,0.0013980508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051416643,0.00060305494,0.0003949531,0.0019664161,0.0008686123,0.0035570778,0.0018864296,0.0012236842,0.012834112],"category_scores_gemma":[0.05379353,0.00050210167,0.0004242534,0.002384162,0.001185718,0.0062878015,0.0036681336,0.0013374513,0.003035146],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008903534,0.0008998432,0.009572206,0.0024187686,0.000050320727,0.0012109106,0.014448342,0.007617961,0.027686365,0.06258612,0.029560968,0.8430579],"study_design_scores_gemma":[0.0007661141,0.0021415828,0.016985489,0.0031812321,0.00019249228,0.0059108087,0.0073216427,0.082328126,0.07724562,0.13237888,0.67112035,0.0004276745],"about_ca_topic_score_codex":0.00077920425,"about_ca_topic_score_gemma":0.0020790414,"teacher_disagreement_score":0.012834112,"about_ca_system_score_codex":0.0006043547,"about_ca_system_score_gemma":0.0015542338,"threshold_uncertainty_score":0.042934358},"labels":[],"label_agreement":null},{"id":"W4389315297","doi":"10.1109/cog57401.2023.10333253","title":"Automatic Bug Detection in Games using LSTM Networks","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Long short term memory; Artificial intelligence; Software bug; Video game; Machine learning; Natural language processing; Programming language; Recurrent neural network; Artificial neural network; Multimedia; Software","score_opus":0.025475635861092215,"score_gpt":0.2833547632484498,"score_spread":0.25787912738735763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389315297","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2009916,0.0012414608,0.7748882,0.000628855,0.00030857755,0.00014881433,0.00080710003,0.018283738,0.0027016979],"genre_scores_gemma":[0.8419794,0.00030155797,0.15445542,0.00022797966,0.00006198922,0.00008645839,0.0006697321,0.0002484502,0.0019689582],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994129,0.00010171203,0.000036678714,0.00021247668,0.00015754579,0.00007871608],"domain_scores_gemma":[0.9989849,0.0004096833,0.0002167395,0.00008067427,0.00024738975,0.000060711678],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000627974,0.0016168184,0.00050032086,0.0015155006,0.0002763067,0.000735679,0.0010791539,0.0009370142,0.0012373887],"category_scores_gemma":[0.0035992414,0.0004484821,0.0006210497,0.0005985507,0.00046878366,0.0015312473,0.00077817094,0.0010442258,0.0003474591],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006824408,0.00036016142,0.015088338,0.00044536305,0.00030431992,0.0010825786,0.00044740984,0.1365721,0.11253511,0.0035261156,0.009556489,0.71939963],"study_design_scores_gemma":[0.000014675851,0.000109165776,0.0033525494,0.000019004592,0.000046790876,0.00011021857,0.000038960352,0.9781286,0.014074549,0.0030775918,0.0010069313,0.000021031008],"about_ca_topic_score_codex":0.009102693,"about_ca_topic_score_gemma":0.014274538,"teacher_disagreement_score":0.009102693,"about_ca_system_score_codex":0.0008476651,"about_ca_system_score_gemma":0.00066897296,"threshold_uncertainty_score":0.018099427},"labels":[],"label_agreement":null},{"id":"W4389337865","doi":"10.1007/s10664-023-10400-0","title":"Bug characterization in machine learning-based systems","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Polytechnique Montréal","funders":"","keywords":"Computer science; Software bug; Software; Component (thermodynamics); Software system; Task (project management); Process (computing); Software engineering; Software development; Focus (optics); Software maintenance; Corrective maintenance; Machine learning; Reliability engineering; Operating system; Systems engineering; Engineering; Preventive maintenance","score_opus":0.024405773824118494,"score_gpt":0.2681081263815844,"score_spread":0.2437023525574659,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389337865","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8627934,0.0010514422,0.13099271,0.00042917317,0.000046078872,0.0001227786,0.0007064349,0.001808361,0.0020496398],"genre_scores_gemma":[0.9758084,0.00008541236,0.023015628,0.000029958948,0.000014797892,0.000027836004,0.0005729179,0.000088673114,0.00035644614],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9933207,0.0019310777,0.00095756777,0.0011000914,0.0022693456,0.0004211822],"domain_scores_gemma":[0.9147225,0.05359252,0.013566099,0.008267517,0.008839646,0.0010117553],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004665167,0.0005314131,0.0005260096,0.0062786965,0.0005330492,0.0014760083,0.000950951,0.0010079887,0.0011892029],"category_scores_gemma":[0.06844034,0.00034198535,0.00055093266,0.0024883032,0.00094554835,0.0031122302,0.0009171902,0.0010315606,0.00024262554],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058314326,0.00058681634,0.58652693,0.00056847994,0.00021044513,0.00041592307,0.0011267341,0.06986167,0.008586907,0.010415968,0.0030204037,0.31809658],"study_design_scores_gemma":[0.000065717686,0.0004946084,0.15054895,0.00021645357,0.0001565505,0.0012336123,0.00058202824,0.80168235,0.011810432,0.031093016,0.0020489672,0.00006732902],"about_ca_topic_score_codex":0.003376165,"about_ca_topic_score_gemma":0.0043997243,"teacher_disagreement_score":0.0062786965,"about_ca_system_score_codex":0.0008380294,"about_ca_system_score_gemma":0.0012247891,"threshold_uncertainty_score":0.024672031},"labels":[],"label_agreement":null},{"id":"W4389518884","doi":"10.18653/v1/2023.findings-emnlp.316","title":"The Vault: A Comprehensive Multilingual Dataset for Advancing Code Understanding and Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Code review; Scripting language; Automatic summarization; Code (set theory); Source code; Artificial intelligence; Natural language processing; KPI-driven code analysis; Vault (architecture); Software quality; Software; Programming language; Software development","score_opus":0.13708468854672312,"score_gpt":0.36369123667287173,"score_spread":0.2266065481261486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518884","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04796427,0.0018385312,0.036483742,0.0015880947,0.00062261126,0.0006153811,0.8572398,0.041460678,0.012187028],"genre_scores_gemma":[0.023481252,0.00021666952,0.0325782,0.0003497596,0.000053820117,0.00056564086,0.93857795,0.0018091527,0.0023676336],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99674654,0.000635577,0.00037936275,0.0009149804,0.0010671985,0.00025626173],"domain_scores_gemma":[0.9930554,0.0019903085,0.000627078,0.0018593135,0.0018918986,0.00057603506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018940215,0.0017385323,0.0006853682,0.00523592,0.0015603646,0.0018125275,0.0030342573,0.0025607941,0.0069988193],"category_scores_gemma":[0.014088125,0.0006138974,0.0015018077,0.0040540933,0.0011000382,0.0030415212,0.0035485413,0.002927563,0.010107615],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059150346,0.00050089206,0.015464321,0.0019381458,0.0001803844,0.0005127649,0.0008220996,0.008376512,0.0068408246,0.0065678963,0.8602309,0.09797367],"study_design_scores_gemma":[0.00060139364,0.00032484086,0.020003682,0.00046570547,0.00011731915,0.0008557518,0.0008244107,0.056226954,0.019387074,0.013210676,0.88770914,0.00027319224],"about_ca_topic_score_codex":0.014860612,"about_ca_topic_score_gemma":0.03305916,"teacher_disagreement_score":0.014860612,"about_ca_system_score_codex":0.0016567862,"about_ca_system_score_gemma":0.002683038,"threshold_uncertainty_score":0.029548228},"labels":[],"label_agreement":null},{"id":"W4389544162","doi":"10.1109/icsme58846.2023.00011","title":"Keynotes","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Code review; Computer science; Software engineering; Software quality; Code (set theory); Source code; Process (computing); Software; Software development; KPI-driven code analysis; Static program analysis; Quality (philosophy); Key (lock); Programming language; Operating system","score_opus":0.027534038281875704,"score_gpt":0.2839848633981639,"score_spread":0.2564508251162882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544162","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00078505673,0.0049041538,0.0033466937,0.17289823,0.3048711,0.00038032985,0.005279559,0.0019184906,0.5056163],"genre_scores_gemma":[0.007954993,0.0024972868,0.0010368177,0.04049586,0.031511936,0.00027962268,0.0022689847,0.0008428616,0.9131117],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99729425,0.00038736517,0.00018762535,0.0004913045,0.0012951929,0.00034437023],"domain_scores_gemma":[0.9877307,0.0026963216,0.0004440436,0.00087269343,0.0059879124,0.0022684503],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0026965584,0.0008785061,0.00071478373,0.0017531059,0.002707232,0.00532722,0.0015168923,0.003992202,0.5834704],"category_scores_gemma":[0.023387892,0.00032819636,0.0006406625,0.001603425,0.00074085675,0.0039595873,0.003063092,0.0049955416,0.3443237],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000023652008,0.000008576113,0.00003870414,0.000037798836,9.4079536e-7,0.000032188706,0.000031318737,0.00001360667,0.000072493895,0.0031116041,0.98663855,0.009990657],"study_design_scores_gemma":[0.000005984123,0.00000680459,0.00010139637,0.000050524402,0.0000015933597,0.000040711486,0.00006741159,0.000025033896,0.0000714851,0.00096716144,0.9986571,0.000004812144],"about_ca_topic_score_codex":0.0044766413,"about_ca_topic_score_gemma":0.0075658434,"teacher_disagreement_score":0.4165296,"about_ca_system_score_codex":0.0029368668,"about_ca_system_score_gemma":0.0028997888,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4389544166","doi":"10.1109/icsme58846.2023.00074","title":"StaticTracker: A Diff Tool for Static Code Warnings","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Commit; Computer science; Static analysis; Code (set theory); Source code; Static program analysis; Software; Software bug; Software quality; Detector; Software engineering; Programming language; Computer security; Operating system; Software development; Database","score_opus":0.02813284726236306,"score_gpt":0.30563111596487097,"score_spread":0.27749826870250793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544166","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041621905,0.00051113317,0.14304544,0.0002635923,0.00026073857,0.00020062187,0.0068419753,0.84232944,0.0023849069],"genre_scores_gemma":[0.11240837,0.0009958283,0.6071234,0.0016310127,0.00032446443,0.0012174126,0.049755435,0.20565726,0.020886771],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99391407,0.00091346947,0.000815145,0.0014602421,0.002473913,0.00042306003],"domain_scores_gemma":[0.9808498,0.008899899,0.0027032506,0.004063809,0.0029358694,0.00054743094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055474727,0.0042285747,0.0016896873,0.0085672485,0.00097623223,0.0023056888,0.005215596,0.0027289612,0.026372017],"category_scores_gemma":[0.032269236,0.0027491224,0.002516494,0.0032882425,0.0012005039,0.006737241,0.0055167587,0.0036655748,0.018777132],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013948923,0.0003168694,0.01626395,0.0024336404,0.00026393158,0.0007206585,0.0010726559,0.005773208,0.014790974,0.007280806,0.38979545,0.559893],"study_design_scores_gemma":[0.0013502981,0.00093688414,0.019097358,0.0012414552,0.0004056393,0.0022735922,0.00053485954,0.30184948,0.12625435,0.03546426,0.50929916,0.0012927214],"about_ca_topic_score_codex":0.0045774826,"about_ca_topic_score_gemma":0.0057779914,"teacher_disagreement_score":0.026372017,"about_ca_system_score_codex":0.00103698,"about_ca_system_score_gemma":0.002943602,"threshold_uncertainty_score":0.08822316},"labels":[],"label_agreement":null},{"id":"W4389544175","doi":"10.1109/icsme58846.2023.00047","title":"Automatic Refactoring Candidate Identification Leveraging Effective Code Representation","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Code refactoring; Computer science; Source code; Identification (biology); Artificial intelligence; Software maintenance; Commit; Machine learning; Code (set theory); Software; Programming language; Software system; Set (abstract data type); Database","score_opus":0.03222432595577605,"score_gpt":0.32703457183795515,"score_spread":0.2948102458821791,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544175","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25559765,0.0008771047,0.73094195,0.00044456625,0.000072555864,0.00021821787,0.00066712225,0.009634985,0.0015458323],"genre_scores_gemma":[0.7012854,0.0002694901,0.29274648,0.00011510201,0.000041719242,0.00013317679,0.0024440594,0.00040732612,0.002557248],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982729,0.00034264138,0.00011630036,0.00046798922,0.0006767018,0.00012353656],"domain_scores_gemma":[0.99313474,0.0022256097,0.0012945539,0.001039514,0.002164049,0.00014146364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013975386,0.00084215304,0.0008169366,0.0034496496,0.00033228574,0.00095075124,0.0010214847,0.00072808686,0.00061180367],"category_scores_gemma":[0.009044488,0.00028548174,0.00049793086,0.0013314981,0.00036065606,0.0016013958,0.00085806457,0.000960056,0.0007358624],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023161898,0.00037137966,0.027297912,0.00020886988,0.00006704907,0.00025591665,0.00030536464,0.021869054,0.061121143,0.0016445427,0.0045041353,0.8821231],"study_design_scores_gemma":[0.000017163586,0.00012966762,0.008052263,0.000043231048,0.000034390658,0.00025291654,0.00007775091,0.9518406,0.033391148,0.0030019246,0.0031329526,0.000025964357],"about_ca_topic_score_codex":0.002643792,"about_ca_topic_score_gemma":0.0051451707,"teacher_disagreement_score":0.0034496496,"about_ca_system_score_codex":0.0005330533,"about_ca_system_score_gemma":0.0012310124,"threshold_uncertainty_score":0.007390976},"labels":[],"label_agreement":null},{"id":"W4389544190","doi":"10.1109/icsme58846.2023.00037","title":"Integrating Visual Aids to Enhance the Code Reviewer Selection Process","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Selection (genetic algorithm); Process (computing); Code (set theory); Programming language; Artificial intelligence","score_opus":0.019851856868040387,"score_gpt":0.36887367494611567,"score_spread":0.3490218180780753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544190","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15086837,0.001412483,0.79551095,0.0026385689,0.0007219268,0.0028399876,0.0010626139,0.032552212,0.012392974],"genre_scores_gemma":[0.23193493,0.000428344,0.7621214,0.00035558,0.00028773842,0.0009629637,0.00052238774,0.00087483064,0.0025119046],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98610455,0.008343192,0.0012544015,0.0015091908,0.002359924,0.00042874922],"domain_scores_gemma":[0.8237861,0.12138268,0.012391739,0.013256945,0.025194865,0.0039876546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022118544,0.0019990925,0.0010140173,0.008114248,0.0010047934,0.005021781,0.0020107029,0.0013490256,0.005370314],"category_scores_gemma":[0.117554896,0.0008217632,0.0007138845,0.0028110943,0.0005987723,0.0042940187,0.0033312934,0.0012612646,0.0022958762],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025215535,0.0008734482,0.014487129,0.0028081408,0.00014097708,0.00060278527,0.0073364708,0.007518637,0.03358112,0.004064789,0.018667009,0.90739805],"study_design_scores_gemma":[0.0024129471,0.0050557535,0.085956536,0.004542387,0.0010522628,0.0045922445,0.0114675155,0.477502,0.12702422,0.044530425,0.23410739,0.0017563533],"about_ca_topic_score_codex":0.0010617337,"about_ca_topic_score_gemma":0.0026543012,"teacher_disagreement_score":0.022118544,"about_ca_system_score_codex":0.0008738813,"about_ca_system_score_gemma":0.0022495184,"threshold_uncertainty_score":0.11697543},"labels":[],"label_agreement":null},{"id":"W4389544206","doi":"10.1109/icsme58846.2023.00060","title":"Finding an Optimal Set of Static Analyzers To Detect Software Vulnerabilities","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Oracle; Set (abstract data type); Vulnerability (computing); Software; Computer security; Static analysis; Data mining; Software engineering; Operating system; Programming language","score_opus":0.0433865343501716,"score_gpt":0.3243299952251804,"score_spread":0.2809434608750088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544206","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5108062,0.001016131,0.46528396,0.0015068479,0.00013525227,0.0006897007,0.0008376904,0.012703474,0.007020678],"genre_scores_gemma":[0.79372686,0.00014996555,0.20337687,0.00030498527,0.000026574855,0.00029816275,0.0006006201,0.00039560994,0.0011203849],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9950354,0.001892027,0.0002693388,0.0013667605,0.0007180311,0.00071844924],"domain_scores_gemma":[0.98801476,0.0071412,0.0011992095,0.002102435,0.00087617716,0.0006662027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004205378,0.001760645,0.001637295,0.0022990152,0.0011881776,0.0015539208,0.0018335127,0.0019715661,0.0031872655],"category_scores_gemma":[0.018079158,0.0010715823,0.001388723,0.0008549641,0.0015901414,0.0031444589,0.0015569993,0.0019330555,0.0013221892],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030282494,0.0020708838,0.03610543,0.0006591933,0.00045423844,0.00049434433,0.00044617377,0.4575631,0.09019345,0.018215002,0.011481232,0.37928864],"study_design_scores_gemma":[0.00026516372,0.0006413477,0.002946283,0.000045180794,0.00011710817,0.00028834623,0.00023449694,0.9603659,0.01367754,0.019319165,0.0020327403,0.00006663281],"about_ca_topic_score_codex":0.002000128,"about_ca_topic_score_gemma":0.0040024524,"teacher_disagreement_score":0.004205378,"about_ca_system_score_codex":0.0016574772,"about_ca_system_score_gemma":0.0038887612,"threshold_uncertainty_score":0.0222404},"labels":[],"label_agreement":null},{"id":"W4389544222","doi":"10.1109/icsme58846.2023.00067","title":"Bugsplainer: Leveraging Code Structures to Explain Software Bugs with Neural Machine Translation","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Dalhousie University","keywords":"Computer science; Debugging; Software bug; Code (set theory); Programming language; Source code; Software; Software engineering; Machine translation; Artificial intelligence; Task (project management); Code smell; Software development; Software quality; Engineering","score_opus":0.03159759519935194,"score_gpt":0.27286243583575276,"score_spread":0.24126484063640083,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544222","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06001404,0.0017573858,0.811882,0.0024737197,0.0006195267,0.00042271285,0.0066868216,0.108922295,0.0072215134],"genre_scores_gemma":[0.3049704,0.0007755452,0.665969,0.00078435784,0.00017028846,0.0004262268,0.016226491,0.0037711773,0.0069064875],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993562,0.00021665452,0.000037217276,0.00023467213,0.00011671327,0.000038558966],"domain_scores_gemma":[0.99633586,0.0026709582,0.00021997548,0.00035816792,0.00034448816,0.0000705448],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00084602623,0.0017833624,0.00037654958,0.0021434778,0.0006323454,0.0009830166,0.0017648544,0.0021303114,0.009075478],"category_scores_gemma":[0.0089997025,0.0006151748,0.0011800475,0.0010315867,0.0006362652,0.0022025187,0.0014367691,0.0020555016,0.003975365],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003707319,0.00039309333,0.0071740276,0.0014704488,0.00016892455,0.0018223438,0.0014449214,0.08814977,0.018915424,0.013489745,0.08050579,0.78609467],"study_design_scores_gemma":[0.00011955482,0.00013058448,0.0011681392,0.00012631468,0.000059153568,0.00042979087,0.00024891904,0.93027925,0.010583484,0.03385833,0.022943357,0.00005312544],"about_ca_topic_score_codex":0.0055061085,"about_ca_topic_score_gemma":0.012700763,"teacher_disagreement_score":0.009075478,"about_ca_system_score_codex":0.0008919745,"about_ca_system_score_gemma":0.0012592959,"threshold_uncertainty_score":0.03036052},"labels":[],"label_agreement":null},{"id":"W4389544307","doi":"10.1109/icsme58846.2023.00013","title":"GPTCloneBench: A comprehensive benchmark of semantic clones and cross-language clones using GPT-3 model and SemanticCloneBench","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Benchmark (surveying); Computer science; Natural language processing; Artificial intelligence; Programming language; Computational biology; Biology; Geography; Cartography","score_opus":0.037414570562637264,"score_gpt":0.32455080197636343,"score_spread":0.2871362314137262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544307","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5758634,0.0040850365,0.104653165,0.0020197004,0.0007690042,0.001361386,0.16311172,0.1157007,0.03243593],"genre_scores_gemma":[0.3117329,0.0009348352,0.1508143,0.0008173997,0.000063938365,0.0011739692,0.5213991,0.007341858,0.00572175],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9944482,0.0009770619,0.0005281436,0.0014881241,0.0021428657,0.0004156931],"domain_scores_gemma":[0.98905784,0.00422078,0.0007156795,0.002802684,0.0026171156,0.0005859333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034227741,0.0028122691,0.0009102265,0.0044536726,0.0011609907,0.0019054752,0.0051580523,0.0023669694,0.0023936522],"category_scores_gemma":[0.016333615,0.0007261923,0.001937186,0.0056166663,0.0014305435,0.0033555352,0.0029707965,0.0022975346,0.001902179],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021295345,0.0024161907,0.062140282,0.0046962597,0.00082275184,0.002312882,0.0013850314,0.2131547,0.027330058,0.013567001,0.3700409,0.30000436],"study_design_scores_gemma":[0.00073760696,0.0017282611,0.03740945,0.00041190287,0.00030904723,0.0019366732,0.0012257749,0.6931381,0.06601067,0.015422706,0.18143217,0.00023766038],"about_ca_topic_score_codex":0.01745941,"about_ca_topic_score_gemma":0.022103854,"teacher_disagreement_score":0.01745941,"about_ca_system_score_codex":0.0023319589,"about_ca_system_score_gemma":0.0027518624,"threshold_uncertainty_score":0.034715593},"labels":[],"label_agreement":null},{"id":"W4389544387","doi":"10.1109/icsme58846.2023.00029","title":"Recommending Code Reviews Leveraging Code Changes with Structured Information Retrieval","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Code (set theory); Information retrieval; Documentation; Source code; Natural language processing; Code review; Class (philosophy); Artificial intelligence; Programming language; Software; Static program analysis; Software development","score_opus":0.054531575323696124,"score_gpt":0.29758076825658425,"score_spread":0.24304919293288813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544387","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28474507,0.01521498,0.6455932,0.002697348,0.0013767616,0.0034306443,0.0042139976,0.028042378,0.014685569],"genre_scores_gemma":[0.4654898,0.0037093298,0.5062328,0.0009878419,0.0007991418,0.00091950805,0.009293528,0.0010041045,0.011563959],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9925309,0.002069704,0.00056362984,0.001354606,0.003293462,0.00018769027],"domain_scores_gemma":[0.9698849,0.011656612,0.003702053,0.0020463718,0.012000389,0.0007097899],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035280322,0.0019229889,0.0012431885,0.009273513,0.0008972077,0.0021546152,0.0011691315,0.0014419704,0.0012393219],"category_scores_gemma":[0.03290772,0.000659592,0.0009822274,0.004024018,0.0004463707,0.0029019173,0.000999198,0.0013062734,0.0022556784],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000528366,0.0004666911,0.019617798,0.0024124044,0.00032474136,0.00076330645,0.0012696823,0.012903518,0.034322757,0.001446683,0.050519556,0.8754245],"study_design_scores_gemma":[0.00036365763,0.0023831574,0.03488064,0.00080872717,0.0010949749,0.0025842506,0.0012810054,0.76740664,0.076362506,0.007293633,0.10508834,0.00045242519],"about_ca_topic_score_codex":0.007205663,"about_ca_topic_score_gemma":0.017796528,"teacher_disagreement_score":0.009273513,"about_ca_system_score_codex":0.0009534182,"about_ca_system_score_gemma":0.0030361183,"threshold_uncertainty_score":0.01865822},"labels":[],"label_agreement":null},{"id":"W4389544451","doi":"10.1109/icsme58846.2023.00018","title":"A Framework for Automating the Measurement of DevOps Research and Assessment (DORA) Metrics","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"DevOps; Computer science; Software engineering; Software deployment","score_opus":0.2633035278652183,"score_gpt":0.4472143890056711,"score_spread":0.1839108611404528,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389544451","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015701895,0.00009426129,0.98543096,0.00031560138,0.000033726938,0.0006633467,0.00076866796,0.0099402005,0.0011831442],"genre_scores_gemma":[0.017494466,0.00009016111,0.9796536,0.00006732694,0.00003080588,0.0008371312,0.0010390874,0.00035339518,0.00043397731],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98196596,0.0048525054,0.003456152,0.0027945405,0.006238291,0.00069252827],"domain_scores_gemma":[0.94256425,0.016578706,0.011321128,0.010885907,0.016716983,0.0019330652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029577404,0.0033565434,0.0018435075,0.019339405,0.0019698688,0.0075419177,0.0028302032,0.0019101669,0.0031873628],"category_scores_gemma":[0.07141388,0.0016836894,0.0026651542,0.008090557,0.0018414126,0.0058709546,0.006060925,0.004937899,0.003180463],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015034873,0.0012019763,0.039587628,0.001509443,0.00040440276,0.0004645018,0.0027229793,0.051890817,0.021127384,0.13998677,0.03343347,0.70752025],"study_design_scores_gemma":[0.00012027816,0.00059598737,0.031069003,0.0010222141,0.00019099469,0.0009581664,0.0013369198,0.6962575,0.024954226,0.12111868,0.12168898,0.00068707875],"about_ca_topic_score_codex":0.016474223,"about_ca_topic_score_gemma":0.018834509,"teacher_disagreement_score":0.029577404,"about_ca_system_score_codex":0.0031567928,"about_ca_system_score_gemma":0.009424147,"threshold_uncertainty_score":0.15642214},"labels":[],"label_agreement":null},{"id":"W4389576314","doi":"10.1109/icsme58846.2023.00038","title":"Breaking the Bento Box: Accelerating Visual Momentum in Data-flow Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Context (archaeology); Reachability; Formative assessment; Data science; World Wide Web; Human–computer interaction; Software engineering; Theoretical computer science","score_opus":0.06596709818660475,"score_gpt":0.34880424756757283,"score_spread":0.2828371493809681,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389576314","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028461546,0.0001785051,0.9577983,0.00029669376,0.00005293738,0.00022508955,0.00013073557,0.010759283,0.0020970241],"genre_scores_gemma":[0.15998843,0.00021064938,0.8349676,0.00020237,0.000023049452,0.00029685322,0.0002463885,0.0021580022,0.0019066972],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9959655,0.0019567742,0.00026473953,0.0006954679,0.00087607803,0.00024134967],"domain_scores_gemma":[0.9742857,0.01899814,0.0010766853,0.00390578,0.0010927126,0.0006409854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008447268,0.0013054829,0.0006311839,0.0020061429,0.00079885137,0.003368536,0.0020496098,0.0011702529,0.007718256],"category_scores_gemma":[0.040452376,0.0010161604,0.00079705496,0.0011223514,0.0025409649,0.008909025,0.007879647,0.0020514152,0.0013630573],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017970832,0.00053992035,0.010742049,0.00080478366,0.00009419418,0.0005141977,0.020176524,0.00893413,0.05525531,0.046462625,0.008190467,0.8464888],"study_design_scores_gemma":[0.0014187626,0.0025535803,0.017774578,0.002054788,0.00031002174,0.001742998,0.008398303,0.402121,0.1324329,0.23217976,0.19826318,0.00075013575],"about_ca_topic_score_codex":0.0021997686,"about_ca_topic_score_gemma":0.0031628218,"teacher_disagreement_score":0.008447268,"about_ca_system_score_codex":0.00059303374,"about_ca_system_score_gemma":0.0016923356,"threshold_uncertainty_score":0.04467398},"labels":[],"label_agreement":null},{"id":"W4389606799","doi":"10.1109/models58315.2023.00037","title":"Automated Domain Modeling with Large Language Models: A Comparative Study","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":72,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Universidad de Murcia","keywords":"Domain (mathematical analysis); Computer science; Set (abstract data type); Modeling language; Domain-specific language; Subject-matter expert; Domain model; Class (philosophy); Domain analysis; Software; Domain knowledge; Natural language processing; Software engineering; Data science; Artificial intelligence; Programming language; Software development; Expert system","score_opus":0.05186317309112483,"score_gpt":0.3344085128492332,"score_spread":0.28254533975810836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389606799","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76060736,0.007834646,0.19962424,0.0024359631,0.00013539674,0.00090596825,0.0021850467,0.0073992526,0.018872099],"genre_scores_gemma":[0.8167381,0.0025894982,0.17366259,0.00026441197,0.00006388422,0.0004189164,0.004470251,0.00068229035,0.0011100128],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9754625,0.017510673,0.0012608723,0.0014405439,0.004029097,0.00029630784],"domain_scores_gemma":[0.7774536,0.19267967,0.004517328,0.016355293,0.008138712,0.00085540605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024707096,0.0012778005,0.0010269373,0.0053912755,0.000836237,0.0035455732,0.002608192,0.0013538225,0.0018509878],"category_scores_gemma":[0.101831324,0.00063272,0.0016725648,0.0053911833,0.0010758353,0.006956459,0.0031370837,0.0017756782,0.0005544955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016765982,0.0021007734,0.0512408,0.002349116,0.0010785419,0.0006751718,0.0076056644,0.22427632,0.0043487716,0.022648737,0.009738463,0.67226106],"study_design_scores_gemma":[0.00030298458,0.0011560558,0.027149353,0.00048297233,0.00058303773,0.0005525776,0.003227594,0.9124775,0.00723313,0.017683871,0.028985187,0.00016579875],"about_ca_topic_score_codex":0.006263711,"about_ca_topic_score_gemma":0.0062338016,"teacher_disagreement_score":0.024707096,"about_ca_system_score_codex":0.0027852484,"about_ca_system_score_gemma":0.002350173,"threshold_uncertainty_score":0.13066518},"labels":[],"label_agreement":null},{"id":"W4389630093","doi":"10.1109/models58315.2023.00029","title":"Automated Grading of Use Cases","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Trent University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Grading (engineering); Computer science; Software; Automation; Sentence; Natural language processing; Software engineering; Artificial intelligence; Matching (statistics); Programming language; Engineering; Mathematics","score_opus":0.06679643628250935,"score_gpt":0.31900178333536433,"score_spread":0.25220534705285497,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389630093","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22399522,0.00057545624,0.7000499,0.000528052,0.00039996346,0.002341904,0.0033558896,0.049169634,0.019584028],"genre_scores_gemma":[0.5187131,0.00028346834,0.4551232,0.00014789404,0.000094005474,0.0006866263,0.009430605,0.0019386837,0.013582502],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9831051,0.0040192087,0.0018507817,0.0025511708,0.0078110397,0.00066269306],"domain_scores_gemma":[0.93733525,0.019454248,0.005211446,0.008215102,0.028610481,0.0011735347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005531958,0.0013829955,0.0013392377,0.0077588614,0.0008157456,0.0033269518,0.0023456304,0.0010653432,0.007174151],"category_scores_gemma":[0.061168566,0.00047655564,0.00085715554,0.0027119364,0.0004323472,0.0021967525,0.0020262317,0.0011241907,0.0043156426],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029122073,0.000442844,0.014078944,0.00041960547,0.0000626123,0.00021924627,0.0005516626,0.0071590096,0.024432648,0.0025162373,0.019196179,0.93062985],"study_design_scores_gemma":[0.00027760895,0.00083993806,0.06884586,0.0004295011,0.00021813581,0.0011406278,0.0014333071,0.6967182,0.13544817,0.016189734,0.07821038,0.0002484631],"about_ca_topic_score_codex":0.0035886064,"about_ca_topic_score_gemma":0.0045344653,"teacher_disagreement_score":0.0077588614,"about_ca_system_score_codex":0.0012477438,"about_ca_system_score_gemma":0.0016798818,"threshold_uncertainty_score":0.029256105},"labels":[],"label_agreement":null},{"id":"W4389822970","doi":"10.1016/j.jss.2023.111935","title":"Incivility detection in open source code review and issue discussions","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Incivility; Open source; Computer science; Tone (literature); Source code; Open source software; Code (set theory); Machine learning; Data science; Artificial intelligence; World Wide Web; Software; Computer security; Psychology; Programming language; Social psychology; Linguistics","score_opus":0.028536473993193052,"score_gpt":0.3096529375676019,"score_spread":0.2811164635744089,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389822970","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.83513707,0.0039997445,0.069776826,0.0102502,0.0017905833,0.0016860256,0.003382971,0.0030995607,0.07087693],"genre_scores_gemma":[0.9545419,0.0005289082,0.027626654,0.0012289513,0.0008850389,0.0007837901,0.0017803187,0.00046807804,0.012156308],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.93719566,0.02244175,0.004940506,0.0060458034,0.026355056,0.0030212111],"domain_scores_gemma":[0.36965767,0.44255018,0.085906595,0.020298583,0.07448781,0.007099098],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04617449,0.00067773304,0.00081855466,0.021928245,0.0033015623,0.0060676476,0.002259744,0.0030564128,0.005059192],"category_scores_gemma":[0.36236864,0.0007204514,0.0008039587,0.0078029237,0.0014310252,0.0076661212,0.0053077424,0.0030667898,0.0013904063],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017897085,0.0006267315,0.44571066,0.002093998,0.00049864146,0.002119989,0.03566585,0.0032871752,0.014803925,0.04496251,0.0564662,0.39197463],"study_design_scores_gemma":[0.00025578807,0.0012285429,0.4663793,0.003873759,0.0010789852,0.0025536532,0.0346334,0.0973574,0.05120428,0.09042519,0.25033978,0.00066992996],"about_ca_topic_score_codex":0.0031325237,"about_ca_topic_score_gemma":0.0044940044,"teacher_disagreement_score":0.95382553,"about_ca_system_score_codex":0.0026934517,"about_ca_system_score_gemma":0.0047685825,"threshold_uncertainty_score":0.24419689},"labels":[],"label_agreement":null},{"id":"W4389980560","doi":"10.1016/j.jss.2023.111934","title":"A survey on machine learning techniques applied to source code","year":2023,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Dalhousie University","funders":"H2020 European Research Council; European Research Council","keywords":"Computer science; Workflow; Source code; Machine learning; Context (archaeology); Software; Software engineering; Artificial intelligence; Task (project management); Code review; Data science; Static program analysis; Software development; Systems engineering; Programming language; Engineering; Database","score_opus":0.02886875562047403,"score_gpt":0.2772920157214117,"score_spread":0.24842326010093768,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389980560","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021907572,0.692631,0.24693824,0.005760988,0.0014912586,0.00046274424,0.0037305423,0.002887691,0.024189983],"genre_scores_gemma":[0.08008909,0.71035737,0.19062188,0.0024478105,0.0017477903,0.00071942096,0.008331302,0.0009401675,0.0047451183],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9917897,0.0019252232,0.0011175362,0.0010808265,0.0038669342,0.00021971039],"domain_scores_gemma":[0.95158273,0.03802421,0.0015408391,0.0022044512,0.006363091,0.0002847948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008219324,0.0011920506,0.0012006619,0.013710264,0.0006927259,0.0026294247,0.002141767,0.0013634118,0.0027595034],"category_scores_gemma":[0.036657948,0.00081198005,0.0020469616,0.019424116,0.00088695734,0.0049877856,0.0012979647,0.0020915163,0.0023226563],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055943914,0.00008331462,0.008838652,0.008642048,0.00019736742,0.00014263873,0.0003165225,0.0036313701,0.0018779321,0.005089032,0.020338818,0.9507864],"study_design_scores_gemma":[0.000047176953,0.00046400234,0.0413124,0.01842663,0.00053979474,0.0023717834,0.0009998383,0.039219942,0.01902314,0.04089694,0.83643705,0.00026138566],"about_ca_topic_score_codex":0.0025652885,"about_ca_topic_score_gemma":0.0025673136,"teacher_disagreement_score":0.013710264,"about_ca_system_score_codex":0.001218177,"about_ca_system_score_gemma":0.0021988594,"threshold_uncertainty_score":0.043468416},"labels":[],"label_agreement":null},{"id":"W4389988541","doi":"10.1109/scam59687.2023.00015","title":"Calibrating Deep Learning-based Code Smell Detection using Human Feedback","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Code smell; Computer science; Code (set theory); Baseline (sea); Context (archaeology); Deep learning; Artificial intelligence; Container (type theory); Software; Software quality; Machine learning; Human–computer interaction; Software development; Engineering; Programming language","score_opus":0.03823697101359752,"score_gpt":0.29761992622234607,"score_spread":0.25938295520874854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389988541","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7890688,0.00093279325,0.1962299,0.00057147816,0.00020413198,0.0002637853,0.0007225129,0.008935903,0.003070633],"genre_scores_gemma":[0.95040566,0.00013750729,0.046514302,0.0002009914,0.00002565467,0.00011456482,0.00094297435,0.00015370785,0.0015046829],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99762887,0.00071213895,0.00015460743,0.0007407206,0.000589001,0.00017465801],"domain_scores_gemma":[0.990704,0.0048646475,0.0010761031,0.00088495127,0.002033183,0.00043705525],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030811515,0.0016215583,0.00070518884,0.00105272,0.00023614167,0.0008191742,0.000972473,0.0010402808,0.0009976166],"category_scores_gemma":[0.018420395,0.00040506432,0.00041530113,0.00045041475,0.0004457539,0.0013946786,0.0010976258,0.0013903109,0.0009290981],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018483696,0.0028576367,0.1326817,0.0011421562,0.00037044467,0.00047063513,0.001154453,0.13292769,0.114348136,0.00061446935,0.010363314,0.60122097],"study_design_scores_gemma":[0.00006006804,0.0007672104,0.02205059,0.00008006174,0.000055265282,0.0001231736,0.00017720833,0.9384587,0.03519015,0.0011171415,0.0018630292,0.000057460446],"about_ca_topic_score_codex":0.0032836539,"about_ca_topic_score_gemma":0.005460295,"teacher_disagreement_score":0.0032836539,"about_ca_system_score_codex":0.0006379155,"about_ca_system_score_gemma":0.0008003245,"threshold_uncertainty_score":0.016294897},"labels":[],"label_agreement":null},{"id":"W4389988565","doi":"10.1109/scam59687.2023.00018","title":"Do Code Quality and Style Issues Differ Across (Non-)Machine Learning Notebooks? Yes!","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Artificial intelligence; Machine learning; Code refactoring; Code (set theory); Software quality; Quality (philosophy); Style (visual arts); Programming style; Programming language; Software; Software development","score_opus":0.05722378713182326,"score_gpt":0.385853588378114,"score_spread":0.32862980124629076,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389988565","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97948945,0.0011604754,0.010055082,0.001388333,0.00011024,0.00015630087,0.001516249,0.0010345245,0.005089407],"genre_scores_gemma":[0.9863927,0.00045799976,0.007360193,0.00055523403,0.00005805359,0.00013404347,0.0020527367,0.000752396,0.0022366298],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9861189,0.0022305944,0.0014351247,0.0027232596,0.006848165,0.00064401637],"domain_scores_gemma":[0.7183767,0.14447184,0.07356199,0.01732215,0.041447636,0.0048197196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013056897,0.00046878768,0.000423927,0.0027437925,0.00062043086,0.0031138803,0.000950927,0.0006840602,0.002499638],"category_scores_gemma":[0.18752803,0.00045254143,0.00059985707,0.0032951713,0.0020964404,0.0044815717,0.0016594668,0.0014537382,0.0010614106],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007342016,0.00024996666,0.79383254,0.00092177495,0.00028403787,0.00030364725,0.010242306,0.0011671208,0.0046433844,0.0019296586,0.007701682,0.17798966],"study_design_scores_gemma":[0.000045440363,0.0003188252,0.9761622,0.0004064515,0.000090667985,0.0004337362,0.004687352,0.0019328018,0.0043147993,0.0020530985,0.009469759,0.000084747604],"about_ca_topic_score_codex":0.0030127931,"about_ca_topic_score_gemma":0.0041376622,"teacher_disagreement_score":0.013056897,"about_ca_system_score_codex":0.0012669795,"about_ca_system_score_gemma":0.0012763689,"threshold_uncertainty_score":0.06905228},"labels":[],"label_agreement":null},{"id":"W4389988584","doi":"10.1109/scam59687.2023.00020","title":"Explaining Transformer-based Code Models: What Do They Learn? When They Do Not Work?","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Calgary","funders":"","keywords":"Computer science; KPI-driven code analysis; Source code; Code (set theory); Code generation; Downstream (manufacturing); Code review; Transformer; Security token; Static program analysis; Artificial neural network; Software engineering; Set (abstract data type); Software; Artificial intelligence; Programming language; Machine learning; Software development; Key (lock); Computer security; Engineering","score_opus":0.06524125755685459,"score_gpt":0.2859045553251602,"score_spread":0.2206632977683056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389988584","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26163107,0.0032259182,0.6848073,0.026324159,0.0002956589,0.0004682098,0.0014393172,0.0062058927,0.015602492],"genre_scores_gemma":[0.77039576,0.0015278526,0.22014463,0.0018798232,0.000096618016,0.0002916401,0.0017823396,0.0006527575,0.0032286672],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99595463,0.0023942809,0.0001425771,0.0006532292,0.0006460112,0.00020932948],"domain_scores_gemma":[0.95603764,0.032105133,0.001592662,0.0065185893,0.0030874559,0.0006584729],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008735626,0.0013904136,0.0005720167,0.0011987734,0.00062069646,0.0025406615,0.0024534317,0.0021706396,0.003991065],"category_scores_gemma":[0.07734156,0.0008211014,0.00094756216,0.00087172445,0.0021444499,0.014653405,0.0018196547,0.004561517,0.0016036065],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042678107,0.00049548986,0.057483807,0.0015886719,0.0004023432,0.0003545665,0.0055250986,0.10507226,0.011653029,0.060267985,0.017880546,0.73884946],"study_design_scores_gemma":[0.00010168532,0.00034633628,0.00814067,0.00044094605,0.00017047123,0.000237813,0.002464759,0.7837906,0.014838618,0.17275105,0.016616134,0.00010092335],"about_ca_topic_score_codex":0.0072084605,"about_ca_topic_score_gemma":0.009859064,"teacher_disagreement_score":0.008735626,"about_ca_system_score_codex":0.0017429796,"about_ca_system_score_gemma":0.0017151093,"threshold_uncertainty_score":0.046198964},"labels":[],"label_agreement":null},{"id":"W4390051453","doi":"10.1145/3638246","title":"Test Generation Strategies for Building Failure Models and Explaining Spurious Failures","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Spurious relationship; Computer science; Test (biology); Machine learning; Surrogate model; Test case; Reliability engineering; Artificial intelligence; Engineering","score_opus":0.13746197668598853,"score_gpt":0.34626461896241467,"score_spread":0.20880264227642614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390051453","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02616278,0.00016228696,0.97008634,0.00026004584,0.00001726051,0.000111910915,0.00024824284,0.0021916097,0.0007595545],"genre_scores_gemma":[0.52967227,0.0001812563,0.46688232,0.00027619879,0.00003051978,0.00038982197,0.0013566695,0.0003894409,0.0008215365],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975466,0.0010613941,0.00016516092,0.00046512485,0.0006240522,0.00013757622],"domain_scores_gemma":[0.98557633,0.010522386,0.0009777109,0.0015159397,0.001222555,0.00018498614],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029061462,0.0019345657,0.00093775586,0.0021804464,0.0004374303,0.0011406628,0.002534576,0.001631492,0.0018448514],"category_scores_gemma":[0.021620609,0.000809661,0.0015713628,0.00087121205,0.0012031484,0.0017023804,0.0014022406,0.0018030617,0.0005477121],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000135098,0.00014970308,0.006780132,0.00014743292,0.00009617845,0.00035349728,0.00018702824,0.88171,0.0043941187,0.010472181,0.001723122,0.09385153],"study_design_scores_gemma":[0.0000102388485,0.000028172992,0.00018129365,0.000011199006,0.000011831283,0.000050078295,0.000011495719,0.991126,0.001534839,0.0067493706,0.00027950853,0.0000059560966],"about_ca_topic_score_codex":0.0042165327,"about_ca_topic_score_gemma":0.0057320944,"teacher_disagreement_score":0.0042165327,"about_ca_system_score_codex":0.0012402211,"about_ca_system_score_gemma":0.0016939944,"threshold_uncertainty_score":0.015369356},"labels":[],"label_agreement":null},{"id":"W4390117375","doi":"10.1109/models-c59198.2023.00101","title":"Towards Understanding and Analyzing Rationale in Commit Messages Using a Knowledge Graph Approach","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal","funders":"Fonds de recherche du Québec – Nature et technologies","keywords":"Commit; Computer science; Pipeline (software); Knowledge graph; Graph; Data science; Component (thermodynamics); Information retrieval; Artificial intelligence; Machine learning; Theoretical computer science; Database; Programming language","score_opus":0.13064983562940144,"score_gpt":0.3264115585535034,"score_spread":0.19576172292410196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390117375","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03686293,0.00077229744,0.94065005,0.0014679906,0.000068486814,0.0005802102,0.008815196,0.006749386,0.0040334957],"genre_scores_gemma":[0.11794749,0.0005666296,0.8578153,0.00023797246,0.00003048936,0.0003093098,0.02061212,0.0005256028,0.0019549942],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99650073,0.0010105892,0.00032757703,0.0007552008,0.0012230526,0.00018281037],"domain_scores_gemma":[0.97984034,0.012781598,0.001667991,0.0024984935,0.0028784606,0.00033310955],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036768892,0.0014902934,0.00075258006,0.016111722,0.001259447,0.004253882,0.0021093343,0.00241109,0.0023339188],"category_scores_gemma":[0.020505492,0.00076409685,0.0017874731,0.0064597796,0.0012470008,0.0068344246,0.002553752,0.0033062117,0.0014748857],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032874913,0.00077926565,0.02679173,0.0020681415,0.00035433008,0.0012140969,0.0040393705,0.063543946,0.020324241,0.062678464,0.027219143,0.79065853],"study_design_scores_gemma":[0.000099086006,0.00021299547,0.013062158,0.0007829574,0.0004191277,0.0008539855,0.004259938,0.63574857,0.027768644,0.21643652,0.10018622,0.00016972834],"about_ca_topic_score_codex":0.017729852,"about_ca_topic_score_gemma":0.046273522,"teacher_disagreement_score":0.017729852,"about_ca_system_score_codex":0.001968193,"about_ca_system_score_gemma":0.0034830435,"threshold_uncertainty_score":0.035253286},"labels":[],"label_agreement":null},{"id":"W4390144898","doi":"10.1007/s10664-023-10431-7","title":"A multi-objective effort-aware approach for early code review prediction and prioritization","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Code review; Code (set theory); Machine learning; Sorting; Prioritization; Empirical research; Source code; Task (project management); Artificial intelligence; Software; Static program analysis; Software engineering; Software development; Management science; Algorithm; Engineering; Systems engineering; Programming language","score_opus":0.039604791623415604,"score_gpt":0.3068962005751302,"score_spread":0.26729140895171455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390144898","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15232903,0.0019231315,0.8350382,0.0010687253,0.00018256412,0.0004923899,0.001295985,0.0038696202,0.0038004213],"genre_scores_gemma":[0.7573309,0.0003081399,0.23747212,0.00021375727,0.00015052732,0.00022489806,0.0012988704,0.00012363786,0.002877246],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973731,0.00052421674,0.0002597634,0.0006592405,0.00091729825,0.00026634583],"domain_scores_gemma":[0.99107236,0.0039726417,0.001427154,0.000443819,0.0025219324,0.0005619943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030515871,0.00180501,0.0020075715,0.0063718157,0.0006515821,0.0019247028,0.0020144188,0.001394081,0.0018232196],"category_scores_gemma":[0.00894935,0.00067981734,0.001107847,0.0029113644,0.00038625376,0.0021121104,0.0015712596,0.0011204819,0.0006317442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059745315,0.0015574874,0.05108267,0.00056010036,0.00055976765,0.00035676238,0.00032760034,0.31763616,0.0113454955,0.0031695885,0.007445448,0.6053615],"study_design_scores_gemma":[0.000015851498,0.00012299277,0.003572707,0.000023112612,0.000063430736,0.000043063454,0.000048244056,0.9926985,0.0011663045,0.001765344,0.00046106393,0.000019410701],"about_ca_topic_score_codex":0.008828976,"about_ca_topic_score_gemma":0.02033119,"teacher_disagreement_score":0.008828976,"about_ca_system_score_codex":0.0009816197,"about_ca_system_score_gemma":0.0031773928,"threshold_uncertainty_score":0.017555177},"labels":[],"label_agreement":null},{"id":"W4390187930","doi":"10.1109/qrs60937.2023.00019","title":"JITBoost: Boosting Just-In-Time Defect Prediction using Boolean Combination of Classifiers","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Boosting (machine learning); Computer science; Machine learning; Artificial intelligence; Pruning; Algorithm; Gradient boosting; Software; Boolean function; Random forest","score_opus":0.044053547415123157,"score_gpt":0.28980026117548835,"score_spread":0.2457467137603652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390187930","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33677348,0.001804095,0.6493241,0.0010439777,0.0003061822,0.00025247416,0.00079615356,0.005405226,0.004294251],"genre_scores_gemma":[0.917412,0.0002247268,0.07863787,0.00036064375,0.00010860075,0.00011489049,0.001387829,0.0001053653,0.0016481271],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99872524,0.00030123643,0.00008166093,0.00030555445,0.0003823721,0.00020388702],"domain_scores_gemma":[0.9955635,0.0021104445,0.00040644145,0.0004298736,0.0012204455,0.0002693299],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026663989,0.0013998409,0.0015190402,0.0017703578,0.0005014863,0.001194086,0.0018359142,0.0011340934,0.0011064818],"category_scores_gemma":[0.007959743,0.00048977026,0.0009515864,0.0010701774,0.000551303,0.001971254,0.0010633629,0.0016253276,0.00051558577],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006217051,0.0005410559,0.038861997,0.00017450232,0.00031464084,0.00020787072,0.00010242509,0.58764863,0.007157393,0.0030161233,0.00763496,0.35371867],"study_design_scores_gemma":[0.000013329087,0.00008319509,0.0007814677,0.0000075998073,0.000025430725,0.000026103035,0.000009195881,0.99616146,0.0009920226,0.0015559462,0.00033710053,0.0000071136747],"about_ca_topic_score_codex":0.0065962933,"about_ca_topic_score_gemma":0.008661999,"teacher_disagreement_score":0.0065962933,"about_ca_system_score_codex":0.0009285565,"about_ca_system_score_gemma":0.0018204525,"threshold_uncertainty_score":0.014101446},"labels":[],"label_agreement":null},{"id":"W4390226062","doi":"10.1145/3631991.3631998","title":"Using Deep Learning and Object-Oriented Metrics to Identify Critical Components in Object-Oriented Systems","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Computer science; Java; Artificial intelligence; Data mining; Suite; Machine learning; Software; Software metric; Deep learning; Principal component analysis; Software system; Software construction; Programming language","score_opus":0.06021645901184579,"score_gpt":0.36079459757355625,"score_spread":0.30057813856171045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390226062","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79616004,0.00080220954,0.19982079,0.0004309263,0.00005648861,0.000060125847,0.0002513824,0.0009794714,0.0014386864],"genre_scores_gemma":[0.97707796,0.00012306248,0.021856118,0.000031567823,0.000011410974,0.000017283955,0.00029579317,0.000018287592,0.00056848564],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993368,0.0002063856,0.000060526007,0.00014217816,0.00016057072,0.000093503],"domain_scores_gemma":[0.9961898,0.0022121326,0.0005384223,0.00023943586,0.00068373635,0.00013647068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001533961,0.0011752313,0.00045490247,0.0018957956,0.00022390878,0.0008186806,0.00060053036,0.000683086,0.00043472645],"category_scores_gemma":[0.007885148,0.00028275212,0.00038867476,0.00092597766,0.0003509269,0.0016687878,0.00066337833,0.0010328537,0.00020286952],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022997672,0.00046100258,0.06441627,0.00013529927,0.00014798398,0.000100186946,0.00020287636,0.62606376,0.0050836354,0.001496144,0.0009626972,0.30070022],"study_design_scores_gemma":[0.000002199925,0.000051898605,0.003542296,0.000008335863,0.0000073546053,0.00001081322,0.000017980858,0.99360895,0.0015287673,0.0011029616,0.00011306169,0.000005378179],"about_ca_topic_score_codex":0.010249739,"about_ca_topic_score_gemma":0.009690526,"teacher_disagreement_score":0.010249739,"about_ca_system_score_codex":0.0009791795,"about_ca_system_score_gemma":0.0008230315,"threshold_uncertainty_score":0.02038014},"labels":[],"label_agreement":null},{"id":"W4390263588","doi":"10.1145/3635439.3635445","title":"Report of the 8th Workshop on Empirical RequirementsEngineering (EmpiRE 2023)","year":2023,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Universitat Politècnica de Catalunya; Khalifa University of Science, Technology and Research; Commonwealth Scientific and Industrial Research Organisation; Hamad Bin Khalifa University; University of Twente; Pennsylvania State University; University of Waterloo; University of Cincinnati; Universiteit Utrecht; University of Pennsylvania","keywords":"Empire; Engineering; Library science; Engineering management; Political science; Computer science; Law","score_opus":0.054257422630188956,"score_gpt":0.3173343984730046,"score_spread":0.26307697584281564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390263588","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0445947,0.03349112,0.2743763,0.21929489,0.0864937,0.0062453886,0.025919275,0.0047993325,0.3047853],"genre_scores_gemma":[0.103938185,0.019987449,0.16698192,0.022934696,0.011144666,0.005783927,0.053251904,0.006556202,0.609421],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98027235,0.009246805,0.0007287242,0.0014356417,0.006885897,0.0014305806],"domain_scores_gemma":[0.963655,0.010862359,0.00093802094,0.0041224537,0.014588965,0.005833196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.046589784,0.0017095185,0.0012574891,0.0024568944,0.0017740526,0.008743218,0.0022590659,0.0042682625,0.062206462],"category_scores_gemma":[0.03284183,0.00085208705,0.0018812813,0.0016442813,0.001026884,0.0049235616,0.0076015163,0.0060589206,0.027705507],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005875535,0.00073168497,0.0011266056,0.0004498729,0.000062056504,0.00032443192,0.0010840879,0.0019239342,0.0044147004,0.009119271,0.80874693,0.17142871],"study_design_scores_gemma":[0.00015079825,0.00027810203,0.002366776,0.00046272634,0.000040719235,0.0001133365,0.0007362249,0.0019097704,0.0038925908,0.0063452153,0.98363924,0.0000645261],"about_ca_topic_score_codex":0.004933875,"about_ca_topic_score_gemma":0.006658749,"teacher_disagreement_score":0.062206462,"about_ca_system_score_codex":0.0026328745,"about_ca_system_score_gemma":0.00957364,"threshold_uncertainty_score":0.24639326},"labels":[],"label_agreement":null},{"id":"W4390411712","doi":"10.1007/s10664-023-10421-9","title":"Using knowledge units of programming languages to recommend reviewers for pull requests: an empirical study","year":2023,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Baseline (sea); Operationalization; Code (set theory); Java; Task (project management); Recommender system; Artificial intelligence; Range (aeronautics); Code review; Information retrieval; Machine learning; Programming language; Natural language processing; Set (abstract data type); Software; Software quality; Software development; Engineering","score_opus":0.15121060831279526,"score_gpt":0.43717184664025377,"score_spread":0.2859612383274585,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390411712","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99503917,0.00020396752,0.0023077922,0.00013427397,0.000010870192,0.000160269,0.00017610771,0.00008822924,0.0018794374],"genre_scores_gemma":[0.99566627,0.00009293976,0.003308662,0.0000550349,0.00001848371,0.00008269205,0.00019712854,0.000033215674,0.0005454827],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97778076,0.011807538,0.0018675673,0.0019828307,0.0058339746,0.00072736805],"domain_scores_gemma":[0.31930277,0.60268307,0.04427809,0.0098335715,0.018287364,0.005615154],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.017069582,0.0005901412,0.0006968158,0.0061153746,0.001292634,0.0044950945,0.0020345014,0.0019931498,0.003370466],"category_scores_gemma":[0.30234155,0.0005761931,0.000648469,0.0049046795,0.0012329669,0.0066303946,0.0015941915,0.0018257987,0.00079263846],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020956043,0.004528131,0.87652963,0.000817513,0.00040525172,0.00052688556,0.010525312,0.0020611912,0.0034155017,0.0008704147,0.0015042669,0.09672028],"study_design_scores_gemma":[0.00046785868,0.003999512,0.9033156,0.00039152315,0.0010880483,0.001487458,0.017254889,0.05477608,0.009449311,0.0032255244,0.0041977996,0.0003462907],"about_ca_topic_score_codex":0.0061583244,"about_ca_topic_score_gemma":0.006708993,"teacher_disagreement_score":0.9829304,"about_ca_system_score_codex":0.0016837134,"about_ca_system_score_gemma":0.0027424777,"threshold_uncertainty_score":0.09027362},"labels":[],"label_agreement":null},{"id":"W4390422940","doi":"10.20508/ijrer.v13i4.14270.g8842","title":"Innovative Cost Estimation for Agile Technology: A Novel Energy Storage Technique Incorporating Modified Planning Poker","year":2023,"lang":"en","type":"article","venue":"International Journal of Renewable Energy Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Agile software development; Computer science; Context (archaeology); Software; Estimation; Workload; Sequence (biology); Fibonacci number; Cost estimate; Scrum; Industrial engineering; Software development; Software engineering; Systems engineering; Data science; Engineering","score_opus":0.071144520269356,"score_gpt":0.37953885543858207,"score_spread":0.3083943351692261,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390422940","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01287665,0.0000627799,0.98549044,0.000045760025,0.000025233687,0.000062329906,0.000029843366,0.00033111358,0.0010757875],"genre_scores_gemma":[0.3260997,0.000109632885,0.67178124,0.000038431215,0.000020582109,0.00017852253,0.00010526684,0.00009485086,0.0015717798],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990984,0.0002181428,0.00005664043,0.0001462459,0.0004356295,0.00004482298],"domain_scores_gemma":[0.997322,0.0014494178,0.000289049,0.00037884602,0.0004995686,0.000061108134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001052542,0.0008787882,0.00058108557,0.0020636113,0.0005262839,0.0009173642,0.0012793039,0.0006417349,0.0024026863],"category_scores_gemma":[0.0058357734,0.00041169405,0.00041967048,0.0016579326,0.00042730014,0.0020472899,0.0009409658,0.0007055716,0.0005109926],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019693718,0.00013248976,0.0030478141,0.00019553736,0.000050788785,0.0002628644,0.00042986617,0.19204172,0.024420233,0.020173462,0.0014584907,0.7575899],"study_design_scores_gemma":[0.000013704297,0.0001447786,0.000975086,0.000025299601,0.000017901035,0.0002556127,0.00008466223,0.974924,0.012108213,0.008884333,0.0025257787,0.000040636995],"about_ca_topic_score_codex":0.0013181224,"about_ca_topic_score_gemma":0.0022174022,"teacher_disagreement_score":0.0024026863,"about_ca_system_score_codex":0.0005076017,"about_ca_system_score_gemma":0.00093698653,"threshold_uncertainty_score":0.008037806},"labels":[],"label_agreement":null},{"id":"W4390502905","doi":"10.1007/s10664-023-10362-3","title":"What is an app store? The software engineering perspective","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of Waterloo","funders":"European Commission","keywords":"App store; Computer science; World Wide Web; Mobile apps; Android (operating system); Software; Download; Operating system","score_opus":0.0255173147224049,"score_gpt":0.30159338532327223,"score_spread":0.27607607060086736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390502905","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10865645,0.023769908,0.11352314,0.19521873,0.00080513826,0.00013022087,0.00084283913,0.00025577348,0.5567978],"genre_scores_gemma":[0.930356,0.017434454,0.021673748,0.0072508105,0.001619965,0.000095497795,0.00031311237,0.00014942892,0.021106988],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.99638605,0.0014447786,0.00019918606,0.00069292285,0.0008213468,0.0004556945],"domain_scores_gemma":[0.9848407,0.009982039,0.00090981415,0.0009550351,0.0020713417,0.001241039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027386541,0.00097096816,0.0010793384,0.0066930284,0.004145586,0.02143852,0.0022405111,0.007065508,0.009774678],"category_scores_gemma":[0.011965342,0.0009070285,0.0006683621,0.0060154703,0.019485373,0.049226996,0.0031228645,0.004868644,0.0017176414],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002189354,0.00008518575,0.0018809543,0.00015062881,0.0000147070405,0.0002674933,0.0020050704,0.00046438043,0.00025640742,0.97710127,0.0031007521,0.014651353],"study_design_scores_gemma":[0.000010366011,0.000047212914,0.0019249951,0.00042499046,0.000036737667,0.00080966594,0.010916125,0.0035086332,0.0006585495,0.9331931,0.04843713,0.000032649616],"about_ca_topic_score_codex":0.011214853,"about_ca_topic_score_gemma":0.0080646,"teacher_disagreement_score":0.02143852,"about_ca_system_score_codex":0.0043008802,"about_ca_system_score_gemma":0.0042220703,"threshold_uncertainty_score":0.032699585},"labels":[],"label_agreement":null},{"id":"W4390605451","doi":"10.1145/3632870","title":"Semantic Code Refactoring for Abstract Data Types","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on Programming Languages","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"National Science Foundation","keywords":"Code refactoring; Computer science; Programming language; Java; Code (set theory); Representation (politics); Equivalence (formal languages); Set (abstract data type); Theoretical computer science; Software","score_opus":0.05658574870714718,"score_gpt":0.34666339177316907,"score_spread":0.2900776430660219,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390605451","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023859601,0.00021425851,0.9617959,0.0002749012,0.00007355337,0.0001393809,0.00025955882,0.011934396,0.0014483947],"genre_scores_gemma":[0.17105468,0.00025769512,0.82158476,0.00029214277,0.000036487298,0.00016392897,0.0012233458,0.0036054237,0.0017815363],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99318415,0.001655924,0.0006189772,0.0010790429,0.0029985625,0.00046340068],"domain_scores_gemma":[0.98397344,0.0062877038,0.0013728547,0.006078336,0.0021348456,0.00015283507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005602782,0.0011690586,0.00080701936,0.0020316096,0.0008307247,0.0018403707,0.0031660965,0.0015660587,0.0026581364],"category_scores_gemma":[0.018038562,0.0008418979,0.0026817115,0.0011835683,0.002959042,0.003902124,0.0034977582,0.003152843,0.00083018234],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006060414,0.00035967055,0.011568079,0.0020661105,0.00034750925,0.0014256237,0.002977626,0.09489054,0.10486075,0.14296728,0.011181267,0.6267495],"study_design_scores_gemma":[0.00017321929,0.0003116155,0.0022098154,0.00046418267,0.00035937288,0.001241181,0.0006194277,0.5051365,0.28072622,0.12748288,0.08108815,0.00018743625],"about_ca_topic_score_codex":0.003035879,"about_ca_topic_score_gemma":0.004248424,"teacher_disagreement_score":0.005602782,"about_ca_system_score_codex":0.0017810974,"about_ca_system_score_gemma":0.0032197759,"threshold_uncertainty_score":0.029630661},"labels":[],"label_agreement":null},{"id":"W4390640589","doi":"10.1016/j.jss.2024.111961","title":"An empirical assessment of different word embedding and deep learning models for bug assignment","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Artificial intelligence; Deep learning; Word embedding; Word2vec; Benchmark (surveying); Natural language processing; Machine learning; Word (group theory); Embedding","score_opus":0.04004862236986989,"score_gpt":0.3478501262090294,"score_spread":0.3078015038391595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390640589","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.944412,0.0018813988,0.048422616,0.0008169377,0.00019350482,0.00009981438,0.0013176989,0.00088412635,0.0019719503],"genre_scores_gemma":[0.9741794,0.0004349013,0.020413242,0.0001061165,0.00007630925,0.00007144082,0.00319695,0.00014095123,0.00138064],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965737,0.001909358,0.00028293376,0.0006907683,0.00036081608,0.00018242492],"domain_scores_gemma":[0.93067414,0.058870554,0.002075369,0.0043357075,0.0032318847,0.000812379],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008404767,0.0016243653,0.0008455399,0.0017752742,0.00055129116,0.0017727315,0.001768808,0.0023251183,0.002451534],"category_scores_gemma":[0.046898894,0.000557774,0.0010197661,0.00159876,0.0011152942,0.007149081,0.001966696,0.0031986723,0.0010765597],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009000932,0.005229869,0.13551746,0.0014758944,0.0013331426,0.00026955237,0.0010599578,0.2608632,0.0058914237,0.0071852803,0.013122594,0.55905074],"study_design_scores_gemma":[0.00019948267,0.0009828545,0.012003145,0.0001310692,0.00032779097,0.00013616879,0.0003040344,0.9752451,0.0020171981,0.007590353,0.0010095676,0.00005324865],"about_ca_topic_score_codex":0.0052537126,"about_ca_topic_score_gemma":0.006200184,"teacher_disagreement_score":0.008404767,"about_ca_system_score_codex":0.001186071,"about_ca_system_score_gemma":0.0010317508,"threshold_uncertainty_score":0.04444915},"labels":[],"label_agreement":null},{"id":"W4390758660","doi":"10.5753/waihcws.2023.233777","title":"Você sabe diferenciar um resumo escrito por humanos do gerado pelo ChatGPT?","year":2023,"lang":"pt","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Humanities; Philosophy; Physics","score_opus":0.04233214720446221,"score_gpt":0.3020221517238567,"score_spread":0.2596900045193945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390758660","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76080793,0.0057431688,0.13485715,0.014718126,0.0011235451,0.00041776366,0.0006779279,0.0033631937,0.078291245],"genre_scores_gemma":[0.9395322,0.0018931372,0.034519535,0.0019450957,0.000277232,0.00011798282,0.00040986657,0.0005379852,0.02076704],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99228936,0.0037445584,0.0004735324,0.0009912715,0.0021207684,0.0003804847],"domain_scores_gemma":[0.94659376,0.031886827,0.005106636,0.0053466475,0.0089659635,0.0021002507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008557084,0.00077289937,0.00058253703,0.002705704,0.002152486,0.006993698,0.001169034,0.0018613285,0.008418043],"category_scores_gemma":[0.074664764,0.0005391905,0.00054095744,0.0017353154,0.00311935,0.007820505,0.0024465513,0.0014292451,0.0031921947],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001126729,0.00030316965,0.117640145,0.0029468248,0.00024099096,0.0020511893,0.3052844,0.0009831189,0.057940148,0.025816701,0.015272397,0.47039413],"study_design_scores_gemma":[0.00008212103,0.0014239473,0.17339478,0.0031893067,0.0006264122,0.006186714,0.26324892,0.012896992,0.039931517,0.03560945,0.4628552,0.00055459485],"about_ca_topic_score_codex":0.004971061,"about_ca_topic_score_gemma":0.0055507515,"teacher_disagreement_score":0.008557084,"about_ca_system_score_codex":0.0012076705,"about_ca_system_score_gemma":0.0018773485,"threshold_uncertainty_score":0.045254707},"labels":[],"label_agreement":null},{"id":"W4390838289","doi":"10.1145/3640331","title":"Method-level Bug Prediction: Problems and Promises","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; York University; University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates","keywords":"Computer science; Software bug; Java; Class (philosophy); Software; Software regression; Granularity; Predictive modelling; Data science; Data mining; Machine learning; Software engineering; Artificial intelligence; Software development; Software quality; Programming language","score_opus":0.13377957027849616,"score_gpt":0.33783300584779025,"score_spread":0.2040534355692941,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390838289","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.109705456,0.08647993,0.6149586,0.15281011,0.0034725927,0.00037032086,0.01050763,0.014015221,0.0076800855],"genre_scores_gemma":[0.560373,0.020079495,0.37502876,0.011103251,0.0052519757,0.0006139664,0.01949944,0.0022209717,0.005829174],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.959857,0.015673386,0.001787013,0.010585613,0.011046789,0.0010501917],"domain_scores_gemma":[0.7298616,0.19185795,0.008964104,0.038232043,0.027086576,0.003997786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.053355675,0.0033034536,0.0036300975,0.005008307,0.0019274217,0.0067918557,0.006931013,0.0048242267,0.0021280614],"category_scores_gemma":[0.17654042,0.0015844858,0.0026470518,0.0069938367,0.0043069473,0.020224754,0.004503762,0.012202567,0.00426308],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005330553,0.0005648086,0.14641513,0.0014599749,0.0006345958,0.00015028534,0.0013514619,0.03893546,0.0018658553,0.012642147,0.0783708,0.71707636],"study_design_scores_gemma":[0.00016373149,0.00076648453,0.07051471,0.0018496772,0.000351188,0.00078442926,0.0026004922,0.6442244,0.005063422,0.18907116,0.084130295,0.00047999332],"about_ca_topic_score_codex":0.02095682,"about_ca_topic_score_gemma":0.012135274,"teacher_disagreement_score":0.053355675,"about_ca_system_score_codex":0.002392705,"about_ca_system_score_gemma":0.0048680203,"threshold_uncertainty_score":0.28217518},"labels":[],"label_agreement":null},{"id":"W4390953316","doi":"10.20944/preprints202401.1133.v1","title":"Interpretable Software Defect Prediction from Project Effort and Static Code Metrics","year":2024,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Fanshawe College","funders":"","keywords":"Interpretability; Computer science; Predictive modelling; Software quality; Machine learning; Reliability (semiconductor); Software; Random forest; Data mining; Software metric; Support vector machine; Software bug; Artificial intelligence; Source code; Code (set theory); Quality (philosophy); Software development; Set (abstract data type); Programming language","score_opus":0.08852903946235081,"score_gpt":0.3473491410098894,"score_spread":0.2588201015475386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390953316","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7948373,0.0002803366,0.19469713,0.0006732189,0.000029816572,0.00013564042,0.0054577067,0.001146884,0.0027419804],"genre_scores_gemma":[0.9754563,0.00007392003,0.02022238,0.00002196924,0.0000091709535,0.00006739077,0.0037137922,0.000038966722,0.0003961943],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99798346,0.0008729675,0.00016056995,0.00039239676,0.00048494912,0.0001056265],"domain_scores_gemma":[0.96974564,0.01985802,0.0043898844,0.0032005142,0.0025551498,0.00025076783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004072748,0.00092655607,0.0003918159,0.0034640422,0.00017850705,0.0011809231,0.00065243186,0.0007808704,0.0012401503],"category_scores_gemma":[0.030162662,0.00022983212,0.00071524567,0.0021078384,0.00038244188,0.0014926796,0.0008497414,0.0009996607,0.000289165],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037412735,0.000449156,0.5053187,0.00042566177,0.00032329842,0.00039729275,0.0011786049,0.28410247,0.0036403765,0.009459658,0.004098221,0.19023249],"study_design_scores_gemma":[0.000022818098,0.00016393862,0.11099159,0.000076774006,0.00006430039,0.000113530004,0.0003196808,0.8683755,0.0025485172,0.015612736,0.0016724657,0.000038177826],"about_ca_topic_score_codex":0.003327494,"about_ca_topic_score_gemma":0.0062085604,"teacher_disagreement_score":0.004072748,"about_ca_system_score_codex":0.00080933125,"about_ca_system_score_gemma":0.0007244332,"threshold_uncertainty_score":0.021538973},"labels":[],"label_agreement":null},{"id":"W4391096329","doi":"10.1109/bigdata59044.2023.10386192","title":"Anaphoric Ambiguity Resolution in Software Requirement Texts","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Ambiguity; Computer science; Baseline (sea); Requirements engineering; Transformer; Artificial intelligence; Software; Architecture; Natural language processing; Software engineering; Engineering; Programming language","score_opus":0.037875549817074056,"score_gpt":0.29708940663528466,"score_spread":0.25921385681821063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391096329","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41196635,0.002429761,0.5676074,0.0019771003,0.00012710667,0.000624528,0.0039096926,0.006593248,0.0047648363],"genre_scores_gemma":[0.76426584,0.00045595734,0.22639246,0.0002626233,0.000069322305,0.00017157057,0.0069881426,0.00011493279,0.0012791867],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99554884,0.002507059,0.00035165626,0.0008221701,0.00065578066,0.00011447182],"domain_scores_gemma":[0.9711963,0.023902955,0.0013547391,0.0013959806,0.0019634787,0.00018658055],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062648896,0.00061877473,0.00064026844,0.0033296312,0.0006382889,0.0012294451,0.0011946657,0.0015255176,0.0012022828],"category_scores_gemma":[0.03325364,0.00035245338,0.00087542355,0.002556998,0.00059294497,0.0033228858,0.0011500178,0.0014035894,0.00074710645],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010323757,0.00081033504,0.019761732,0.0015621956,0.00015186255,0.00094479846,0.0018098124,0.30976626,0.01596311,0.013146966,0.01586239,0.6191882],"study_design_scores_gemma":[0.000051318595,0.00007324646,0.0029163521,0.00003408412,0.000024833595,0.00020848968,0.00023008176,0.9778205,0.0082065165,0.007501815,0.0029118434,0.00002096088],"about_ca_topic_score_codex":0.005024705,"about_ca_topic_score_gemma":0.0043364586,"teacher_disagreement_score":0.0062648896,"about_ca_system_score_codex":0.0013582789,"about_ca_system_score_gemma":0.0009044997,"threshold_uncertainty_score":0.033132315},"labels":[],"label_agreement":null},{"id":"W4391125455","doi":"10.48550/arxiv.2401.10359","title":"Keeping Deep Learning Models in Check: A History-Based Approach to Mitigate Overfitting","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); University of Alberta","funders":"","keywords":"Overfitting; Computer science; Machine learning; Artificial intelligence; Classifier (UML); Early stopping; Deep learning; Artificial neural network","score_opus":0.09718714126425883,"score_gpt":0.19490207643199758,"score_spread":0.09771493516773876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391125455","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11409079,0.0016462922,0.8707483,0.0014105656,0.00021048094,0.00017550151,0.00029334464,0.008693683,0.0027310192],"genre_scores_gemma":[0.8569569,0.0004869051,0.13503361,0.0010887005,0.00016625869,0.00018369143,0.00088815717,0.0006816443,0.0045141764],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973652,0.0004816803,0.00025182983,0.0007103626,0.0008971967,0.0002937525],"domain_scores_gemma":[0.98887664,0.0048539913,0.0016994152,0.0019900908,0.0021100629,0.00046973152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053072413,0.0024461115,0.0020760808,0.002103285,0.00106427,0.0019098222,0.003806468,0.002298218,0.00204635],"category_scores_gemma":[0.022382313,0.0012336029,0.0014665918,0.0010627112,0.00129387,0.00387464,0.0028077275,0.0047600586,0.00094109494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004924881,0.00062538777,0.025169147,0.0002353319,0.0003903384,0.00055615156,0.00050193176,0.4686149,0.011180298,0.0034427012,0.008159235,0.48063213],"study_design_scores_gemma":[0.00001216348,0.00013992732,0.0012515116,0.000041031748,0.00006103878,0.00008714507,0.000030274712,0.9903172,0.004314002,0.002702127,0.0010175843,0.000025996585],"about_ca_topic_score_codex":0.010251684,"about_ca_topic_score_gemma":0.014031635,"teacher_disagreement_score":0.010251684,"about_ca_system_score_codex":0.0014835685,"about_ca_system_score_gemma":0.0024268113,"threshold_uncertainty_score":0.028067708},"labels":[],"label_agreement":null},{"id":"W4391164126","doi":"10.1109/tse.2024.3358258","title":"Multi-Language Software Development: Issues, Challenges, and Solutions","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"Office of Naval Research","keywords":"Computer science; Interoperability; World Wide Web; Software development; Software; Software engineering; Programming language; Data science","score_opus":0.03412181408304442,"score_gpt":0.26704306635803715,"score_spread":0.23292125227499272,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391164126","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48844504,0.04794777,0.08206375,0.3463203,0.0017764547,0.00044683454,0.00033447138,0.0013531366,0.031312186],"genre_scores_gemma":[0.8720332,0.02241936,0.084394455,0.009981507,0.0015052917,0.00036875517,0.00042975848,0.00049948605,0.0083682],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9692848,0.015427383,0.002109468,0.0027778733,0.007937842,0.0024626993],"domain_scores_gemma":[0.8963726,0.068931736,0.010060695,0.0037398448,0.014714793,0.006180368],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026497122,0.00077042496,0.00055551913,0.0050166207,0.0065710032,0.0108392015,0.0027634548,0.0035716933,0.002402509],"category_scores_gemma":[0.05956748,0.0009828911,0.00075439236,0.0063553387,0.0059029735,0.020661734,0.008221827,0.0046926783,0.0008730162],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015364806,0.0005631735,0.08236866,0.0033042473,0.000083000035,0.0034753967,0.13550325,0.0016959131,0.0073132208,0.05162569,0.032668024,0.6812458],"study_design_scores_gemma":[0.000052511667,0.00042131782,0.0644218,0.0036629206,0.00009867783,0.0070894645,0.505407,0.014295315,0.0053298334,0.094018824,0.30479047,0.00041182776],"about_ca_topic_score_codex":0.005019992,"about_ca_topic_score_gemma":0.008553477,"teacher_disagreement_score":0.026497122,"about_ca_system_score_codex":0.0039890525,"about_ca_system_score_gemma":0.00853379,"threshold_uncertainty_score":0.14013189},"labels":[],"label_agreement":null},{"id":"W4391164257","doi":"10.1109/tse.2024.3358283","title":"Tracking the Evolution of Static Code Warnings: The State-of-the-Art and a Better Approach","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Static analysis; Static program analysis; Code (set theory); Workflow; Tracking (education); Source code; Software engineering; Software; Tracking system; Software evolution; Code smell; Programming language; Software development; Artificial intelligence; Software quality; Database; Software construction","score_opus":0.011706816219258082,"score_gpt":0.22705450869958402,"score_spread":0.21534769248032593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391164257","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51566744,0.050157752,0.34230635,0.010305639,0.0018554458,0.0005777822,0.020328576,0.05375492,0.0050461534],"genre_scores_gemma":[0.61111146,0.0077378056,0.33123764,0.0016252139,0.00066558406,0.00040362866,0.040725075,0.0027338983,0.003759638],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9867621,0.0025923366,0.0015778723,0.004425033,0.0039596627,0.0006828855],"domain_scores_gemma":[0.9419053,0.024331434,0.0074346154,0.013987126,0.010704978,0.0016363816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075858305,0.0029000607,0.0024506582,0.014825568,0.0016573872,0.004984097,0.0041628825,0.003548477,0.0009105453],"category_scores_gemma":[0.045119207,0.0013068726,0.0020152503,0.00989883,0.0014553602,0.0077514243,0.0032565317,0.004069152,0.0011963693],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006761364,0.00075125083,0.20980676,0.0031569896,0.0007718878,0.0005244736,0.0017905574,0.02175712,0.015652578,0.0035150228,0.03995555,0.7016417],"study_design_scores_gemma":[0.0003058667,0.00096752256,0.17309283,0.0015234207,0.0011902315,0.0021000372,0.002961935,0.64536315,0.028385403,0.019153725,0.124377295,0.0005786522],"about_ca_topic_score_codex":0.023635311,"about_ca_topic_score_gemma":0.030410897,"teacher_disagreement_score":0.023635311,"about_ca_system_score_codex":0.0013403613,"about_ca_system_score_gemma":0.0036255557,"threshold_uncertainty_score":0.04699546},"labels":[],"label_agreement":null},{"id":"W4391165725","doi":"10.1109/access.2024.3358201","title":"Software Defect Prediction Using an Intelligent Ensemble-Based Model","year":2024,"lang":"en","type":"article","venue":"IEEE Access","topic":"Software Engineering Research","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"Princess Nourah Bint Abdulrahman University","keywords":"Computer science; Software bug; Software; Artificial intelligence; Operating system","score_opus":0.09921575589353612,"score_gpt":0.35804787843194325,"score_spread":0.25883212253840715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391165725","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24144478,0.0008439923,0.75233835,0.00038416203,0.00010106993,0.000066300345,0.0003789662,0.0015549635,0.0028874513],"genre_scores_gemma":[0.9505161,0.00032606735,0.047081973,0.00006993588,0.00004933742,0.0000735904,0.00038568766,0.000027012411,0.0014703041],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99956125,0.00007714351,0.00003048797,0.00013459996,0.00014075663,0.000055734316],"domain_scores_gemma":[0.99918586,0.000333574,0.000104387516,0.00006483781,0.00028007466,0.000031371706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00091495056,0.0007441463,0.0011388893,0.001272867,0.000324503,0.00073037663,0.0010572738,0.0006971688,0.0006851896],"category_scores_gemma":[0.0018375288,0.00032650333,0.00091904745,0.0008106005,0.00016308234,0.001069273,0.00050738844,0.00068891887,0.00024670895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009638535,0.00011793871,0.009530941,0.0000303909,0.00014443732,0.00009200962,0.00005228106,0.87641114,0.0019748504,0.00086972635,0.0010626598,0.10961721],"study_design_scores_gemma":[0.0000010578008,0.000010203702,0.00028416168,0.0000014280834,0.000007649227,0.0000062323857,0.0000017576904,0.99933535,0.00010928851,0.00019159504,0.00004938804,0.0000018910229],"about_ca_topic_score_codex":0.01088045,"about_ca_topic_score_gemma":0.009946673,"teacher_disagreement_score":0.01088045,"about_ca_system_score_codex":0.00047723725,"about_ca_system_score_gemma":0.0005179264,"threshold_uncertainty_score":0.021634221},"labels":[],"label_agreement":null},{"id":"W4391179708","doi":"10.1016/j.jss.2024.111974","title":"API usage templates via structural generalization","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Template; Generalization; Computer science; Programming language; Mathematics","score_opus":0.014502280135902434,"score_gpt":0.2619418544786123,"score_spread":0.24743957434270986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391179708","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.067611486,0.00058654044,0.90852237,0.00085843384,0.00006378003,0.00065713725,0.0027935538,0.010890223,0.0080164345],"genre_scores_gemma":[0.2502957,0.00044968657,0.73802567,0.00034951058,0.000042937205,0.0007759053,0.0058040046,0.0016610952,0.002595469],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9934989,0.0020079496,0.00087593345,0.0015479999,0.0018713409,0.00019777102],"domain_scores_gemma":[0.96840906,0.015685875,0.0032994016,0.009647898,0.0026277492,0.00033009297],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041859653,0.0009903952,0.0006823871,0.0034126653,0.00084947713,0.0023578461,0.0016550993,0.0012326157,0.0032072314],"category_scores_gemma":[0.041463334,0.0010789427,0.0016550637,0.0036070975,0.0017183962,0.006107818,0.0029601494,0.0018271862,0.0016957257],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003221883,0.0003559144,0.039437804,0.0010049713,0.0001485218,0.00065097545,0.003650844,0.023906173,0.011333164,0.06374431,0.013688739,0.84175646],"study_design_scores_gemma":[0.0001200152,0.0005130212,0.023478143,0.0009665424,0.0002625848,0.0035811597,0.0021866704,0.46078882,0.052594215,0.30065402,0.15460028,0.00025443835],"about_ca_topic_score_codex":0.0023099752,"about_ca_topic_score_gemma":0.0043585007,"teacher_disagreement_score":0.0041859653,"about_ca_system_score_codex":0.0008286625,"about_ca_system_score_gemma":0.0018936809,"threshold_uncertainty_score":0.022137761},"labels":[],"label_agreement":null},{"id":"W4391182719","doi":"10.1134/s0361768823080157","title":"A Metrics Suite for Measuring Indirect Coupling Complexity","year":2023,"lang":"en","type":"article","venue":"Programming and Computer Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Suite; Computer science; Coupling (piping); Engineering; Geography; Mechanical engineering","score_opus":0.08919272441544598,"score_gpt":0.29219110478311233,"score_spread":0.20299838036766635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391182719","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41479355,0.001698599,0.5591984,0.0005654779,0.00014645889,0.0015763905,0.0053310497,0.004898101,0.011791922],"genre_scores_gemma":[0.68864715,0.00037063568,0.30128756,0.00008144356,0.000055710447,0.0013778104,0.006487609,0.00038342085,0.0013086321],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9720538,0.008316573,0.0033639953,0.0013878692,0.014402353,0.00047536698],"domain_scores_gemma":[0.90544015,0.03851768,0.019717814,0.0078105703,0.026391404,0.0021223507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011623611,0.0021263235,0.0010575535,0.018087337,0.00094928266,0.0032151912,0.0014681212,0.00096306595,0.0014686314],"category_scores_gemma":[0.07473597,0.0004476042,0.00085807,0.010797063,0.0009225137,0.0037174253,0.003038706,0.0011392848,0.00049234985],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044380265,0.0010394683,0.27770448,0.0015576325,0.0011602563,0.00029991355,0.0021327755,0.06968607,0.023039067,0.023274895,0.01281116,0.58685046],"study_design_scores_gemma":[0.00017956004,0.004016369,0.23325533,0.00077901874,0.00061234756,0.00095114065,0.0016364343,0.64314735,0.038656306,0.044455204,0.031894654,0.00041639927],"about_ca_topic_score_codex":0.0021417518,"about_ca_topic_score_gemma":0.002699014,"teacher_disagreement_score":0.018087337,"about_ca_system_score_codex":0.0012233084,"about_ca_system_score_gemma":0.001999239,"threshold_uncertainty_score":0.061472237},"labels":[],"label_agreement":null},{"id":"W4391395450","doi":"10.1007/s10515-024-00413-4","title":"An extensive study of the effects of different deep learning models on code vulnerability detection in Python code","year":2024,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Natural Science Foundation of China","keywords":"Python (programming language); Computer science; Artificial intelligence; Deep learning; Word2vec; Machine learning; Programming language","score_opus":0.012330757009532687,"score_gpt":0.2627646923285514,"score_spread":0.2504339353190187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391395450","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9795869,0.00092242926,0.016386446,0.00029978697,0.00006391457,0.000025149568,0.00036500377,0.001100591,0.0012496974],"genre_scores_gemma":[0.988005,0.00024129079,0.010071548,0.000083730694,0.00001744812,0.0000148149775,0.00064944,0.00009165767,0.0008250257],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988477,0.0003154883,0.00006801059,0.0002630634,0.00032610729,0.00017952014],"domain_scores_gemma":[0.984554,0.011230401,0.0009483765,0.0014511566,0.0015083216,0.00030782612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015522268,0.00090216467,0.00053455273,0.0007822523,0.00044734296,0.0005041836,0.00079464447,0.000777821,0.00080270786],"category_scores_gemma":[0.01580907,0.00025404713,0.0005652813,0.000715752,0.00065353053,0.0016328029,0.00078652287,0.001403622,0.00021065214],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002216878,0.0015200431,0.065496996,0.00092811213,0.00049496145,0.0005280684,0.0004509055,0.55909485,0.03589426,0.0027846098,0.011673903,0.3189164],"study_design_scores_gemma":[0.000026798849,0.0003524541,0.010475862,0.000040589548,0.0000845883,0.00009264514,0.00010219442,0.96882623,0.017381521,0.0017570239,0.0008364712,0.000023597928],"about_ca_topic_score_codex":0.0103052985,"about_ca_topic_score_gemma":0.012504211,"teacher_disagreement_score":0.0103052985,"about_ca_system_score_codex":0.000847949,"about_ca_system_score_gemma":0.0011519601,"threshold_uncertainty_score":0.020490587},"labels":[],"label_agreement":null},{"id":"W4391467874","doi":"10.1007/s11334-024-00550-9","title":"Prioritizing unit tests using object-oriented metrics, centrality measures, and machine learning algorithms","year":2024,"lang":"en","type":"article","venue":"Innovations in Systems and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Centrality; Machine learning; Algorithm; Unit (ring theory); Data mining; Mathematics; Statistics; Mathematics education","score_opus":0.045021694363796806,"score_gpt":0.2854496396287489,"score_spread":0.24042794526495212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391467874","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20514518,0.0017931262,0.78238946,0.0009862641,0.00026325355,0.00046942488,0.00037951992,0.002558289,0.0060155042],"genre_scores_gemma":[0.67040133,0.0003264792,0.3254839,0.00015331287,0.00018074785,0.00020049022,0.00086019543,0.0004942382,0.0018991395],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98633873,0.0048688194,0.0011826195,0.0012327345,0.005617817,0.0007593252],"domain_scores_gemma":[0.92074263,0.05443828,0.006263127,0.003632547,0.012035838,0.0028875766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010756467,0.0018916973,0.002345692,0.012227617,0.0010209,0.003913245,0.0031066844,0.0015016291,0.0031305011],"category_scores_gemma":[0.0727488,0.00047663628,0.0008658395,0.0048692445,0.0010728622,0.0044995993,0.0022200996,0.0014985144,0.00074497785],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012703486,0.0007632282,0.07779477,0.0006741686,0.0003612786,0.00041607703,0.0003271949,0.10501447,0.010325724,0.036544226,0.0064039784,0.76010454],"study_design_scores_gemma":[0.00014348538,0.0006423581,0.010553215,0.00013452748,0.00017450607,0.00032938545,0.00025511478,0.89905006,0.012560987,0.07288476,0.0032002516,0.00007137795],"about_ca_topic_score_codex":0.0055227135,"about_ca_topic_score_gemma":0.0074879834,"teacher_disagreement_score":0.012227617,"about_ca_system_score_codex":0.0017039104,"about_ca_system_score_gemma":0.004575239,"threshold_uncertainty_score":0.056886315},"labels":[],"label_agreement":null},{"id":"W4391506256","doi":"10.1145/3613905.3650867","title":"To Search or To Gen? Exploring the Synergy between Generative AI and Web Search in Programming","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Generative grammar; Process (computing); Artificial intelligence; Human–computer interaction; Data science; World Wide Web; Programming language","score_opus":0.11397345323612454,"score_gpt":0.3560189740201909,"score_spread":0.24204552078406635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391506256","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9171005,0.00031805824,0.060966123,0.0019330427,0.0000056991944,0.00006524206,0.000015943942,0.00009339211,0.01950193],"genre_scores_gemma":[0.9839759,0.00010443984,0.015270099,0.00007891702,0.0000018861952,0.00002542942,0.000010067236,0.000024627769,0.00050857215],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99522287,0.0037860011,0.000079830774,0.00025110826,0.0003961005,0.00026403426],"domain_scores_gemma":[0.9616743,0.03520502,0.0011302123,0.00093931495,0.0005664624,0.00048466792],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005563189,0.00035596982,0.00030111775,0.0017034075,0.0014000889,0.0047422573,0.0010302298,0.0009889341,0.0015283361],"category_scores_gemma":[0.020966917,0.0005432604,0.00029611995,0.0014808938,0.006831751,0.007804265,0.0031316814,0.0016500085,0.00014378004],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043038593,0.00052623375,0.058492087,0.0006906132,0.0000775574,0.0021750857,0.5529115,0.008271542,0.010194962,0.21283004,0.00083722384,0.15256274],"study_design_scores_gemma":[0.00013266988,0.0006140733,0.04601405,0.00055194774,0.00012835552,0.0029938933,0.4414259,0.11772171,0.01058126,0.35442522,0.02526391,0.00014700883],"about_ca_topic_score_codex":0.0022681863,"about_ca_topic_score_gemma":0.0036076545,"teacher_disagreement_score":0.005563189,"about_ca_system_score_codex":0.0012198839,"about_ca_system_score_gemma":0.0019463476,"threshold_uncertainty_score":0.02942127},"labels":[],"label_agreement":null},{"id":"W4391559702","doi":"10.1109/tse.2024.3363223","title":"DynAMICS: A Tool-Based Method for the Specification and Dynamic Detection of Android Behavioral Code Smells","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code smell; Computer science; Android (operating system); Source code; Software; Code (set theory); Software quality; Programming language; Artificial intelligence; Software engineering; Software development; Operating system","score_opus":0.01749395245963264,"score_gpt":0.29060479010718,"score_spread":0.27311083764754734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391559702","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016527683,0.000068964786,0.8620592,0.00013426504,0.000098668665,0.0004782916,0.0010765805,0.13200876,0.0024225938],"genre_scores_gemma":[0.03772072,0.00025222445,0.9157574,0.00034084104,0.00008638324,0.0023811802,0.003806286,0.028512849,0.011142113],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99558216,0.0007017394,0.00057117187,0.0009919972,0.0018769763,0.00027590198],"domain_scores_gemma":[0.99163955,0.004374263,0.00081732584,0.0015381862,0.001311829,0.00031889515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036611003,0.0036963602,0.0014463842,0.0042786594,0.00095170294,0.0032398615,0.0028842664,0.0025126452,0.015251922],"category_scores_gemma":[0.01655347,0.002184758,0.0026538824,0.0010223867,0.0016397358,0.0041125044,0.004640372,0.002968779,0.010305689],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011403526,0.0005676309,0.011449877,0.002269701,0.00043227826,0.0022786693,0.0036621806,0.010582208,0.095589496,0.050983593,0.12337408,0.6976699],"study_design_scores_gemma":[0.0006552626,0.0004670155,0.0057257684,0.00080203486,0.00028032355,0.002553538,0.0007723587,0.3009022,0.14638427,0.033086907,0.507548,0.00082236127],"about_ca_topic_score_codex":0.0031696847,"about_ca_topic_score_gemma":0.004216821,"teacher_disagreement_score":0.015251922,"about_ca_system_score_codex":0.0010398964,"about_ca_system_score_gemma":0.0034405622,"threshold_uncertainty_score":0.051022828},"labels":[],"label_agreement":null},{"id":"W4391562352","doi":"10.18260/1-2--40663","title":"Work In Progress: CodeCapture: A Tool to Attain Insight into the Programming Development Process","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Process (computing); Computer science; Work in process; Work (physics); Process management; Software engineering; Programming language; Engineering; Operations management; Mechanical engineering","score_opus":0.015880794509510302,"score_gpt":0.29278173205949437,"score_spread":0.27690093754998407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391562352","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037469566,0.0008841793,0.7383979,0.0014892442,0.00077181665,0.00089689967,0.005393238,0.18882282,0.025874363],"genre_scores_gemma":[0.1698414,0.00080848555,0.7665016,0.0004049809,0.00026367445,0.0013075843,0.013566054,0.024608042,0.022698272],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9942264,0.0014267034,0.00046625896,0.0010450764,0.002559827,0.00027575862],"domain_scores_gemma":[0.9672784,0.01726127,0.0023092723,0.005585208,0.0058165994,0.0017491301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072140554,0.0020678795,0.0006698808,0.005996011,0.0011063241,0.0040397653,0.0023501003,0.0014950302,0.01487704],"category_scores_gemma":[0.03660769,0.0009928987,0.00078983884,0.0029784855,0.0010159174,0.006180761,0.00382428,0.0025378764,0.0072087157],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004729061,0.00059450616,0.006185132,0.0012289048,0.00009870809,0.0004903928,0.0063214186,0.0061140615,0.016001035,0.008112552,0.081315845,0.8730646],"study_design_scores_gemma":[0.0004887793,0.0017184203,0.017204916,0.0015734734,0.00026423394,0.002172143,0.0023118008,0.12754835,0.09484469,0.021268213,0.7299232,0.00068183494],"about_ca_topic_score_codex":0.0017349999,"about_ca_topic_score_gemma":0.0018130747,"teacher_disagreement_score":0.01487704,"about_ca_system_score_codex":0.0007402636,"about_ca_system_score_gemma":0.0026565327,"threshold_uncertainty_score":0.049768686},"labels":[],"label_agreement":null},{"id":"W4391579662","doi":"10.1145/3597503.3623321","title":"FuzzSlice: Pruning False Positives in Static Analysis Warnings through Function-Level Fuzzing","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Fuzz testing; False positive paradox; Static analysis; Computer science; Function (biology); Pruning; Crash; Code (set theory); False positives and false negatives; Data mining; Machine learning; Software; Programming language; Set (abstract data type)","score_opus":0.040309372114255226,"score_gpt":0.31567045045888265,"score_spread":0.2753610783446274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391579662","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.061637793,0.0010994697,0.84512186,0.00075665425,0.0002878223,0.00029096103,0.0008217047,0.08591338,0.0040703453],"genre_scores_gemma":[0.39317507,0.0005804811,0.58856434,0.0008201631,0.00014264196,0.00027110835,0.0023505026,0.006791229,0.007304494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99514955,0.0010755324,0.00033589624,0.0009022799,0.0021181072,0.00041861986],"domain_scores_gemma":[0.98657155,0.0071336185,0.00089948747,0.003811135,0.0013180166,0.00026611047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00409829,0.002160077,0.0014301832,0.0037533394,0.0010133856,0.0024802526,0.0029708913,0.0021682417,0.0062586614],"category_scores_gemma":[0.02048699,0.0012255937,0.0016033304,0.0014552473,0.002036362,0.0047149397,0.003479988,0.0022677656,0.0025802634],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017272944,0.00039532335,0.014989166,0.00091872405,0.00041252037,0.0013660076,0.0012154483,0.050209615,0.08919828,0.023850853,0.030473422,0.7852433],"study_design_scores_gemma":[0.00027751416,0.0005348511,0.0049917907,0.000344999,0.00039647706,0.0014486446,0.0002553106,0.74855953,0.16747786,0.05075706,0.024745777,0.00021018446],"about_ca_topic_score_codex":0.0029195258,"about_ca_topic_score_gemma":0.0045978795,"teacher_disagreement_score":0.0062586614,"about_ca_system_score_codex":0.00085405377,"about_ca_system_score_gemma":0.0016440621,"threshold_uncertainty_score":0.021674097},"labels":[],"label_agreement":null},{"id":"W4391649227","doi":"10.1002/smr.2657","title":"Practitioners' expectations on automated release note generation techniques","year":2024,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Software release life cycle; Task (project management); Software engineering; Software; Data science; Software development; Process management; World Wide Web; Engineering; Software construction; Systems engineering","score_opus":0.015397722247283613,"score_gpt":0.31779093120839425,"score_spread":0.30239320896111066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391649227","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9085012,0.0018049349,0.053466603,0.02256446,0.00021597577,0.0005695226,0.0002001725,0.0011871522,0.011489964],"genre_scores_gemma":[0.95673394,0.0013202894,0.03675858,0.0022757004,0.00009813574,0.0005136821,0.00022341375,0.00021002756,0.0018661951],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.8901858,0.057476804,0.009026631,0.0048268936,0.0345254,0.003958456],"domain_scores_gemma":[0.37022746,0.4114439,0.051323876,0.02451642,0.12865649,0.013831865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.12054744,0.000586592,0.00048255594,0.0035976921,0.0015107039,0.0070787347,0.0028018008,0.003765097,0.003067058],"category_scores_gemma":[0.38580737,0.0009372883,0.0008621226,0.0016456626,0.0019737114,0.0069096275,0.0030386627,0.00358277,0.0019112389],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009766447,0.0015719412,0.18964805,0.0043507973,0.00014850784,0.0031410593,0.32471335,0.004260749,0.038435206,0.005969079,0.01596022,0.41082442],"study_design_scores_gemma":[0.0004142115,0.007821245,0.2343015,0.011087993,0.00045198362,0.0074278126,0.5094454,0.04426487,0.023926627,0.013566921,0.14606461,0.0012268586],"about_ca_topic_score_codex":0.0019440933,"about_ca_topic_score_gemma":0.001684653,"teacher_disagreement_score":0.12054744,"about_ca_system_score_codex":0.003621503,"about_ca_system_score_gemma":0.0047377,"threshold_uncertainty_score":0.6375234},"labels":[],"label_agreement":null},{"id":"W4391660146","doi":"10.1007/978-3-031-53227-6_10","title":"Understanding User Feedback in Software Ecosystems: A Study on Challenges and Mitigation Strategies","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Software; Thematic analysis; User story; User requirements document; User experience design; Human–computer interaction; Data science; Software development; Software engineering; Qualitative research","score_opus":0.05618622903845838,"score_gpt":0.2710712318472808,"score_spread":0.21488500280882245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391660146","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9941345,0.00028458654,0.0024510175,0.00096091704,0.000014277159,0.000058801055,0.000044149583,0.000028637618,0.0020233733],"genre_scores_gemma":[0.99759585,0.00013257653,0.0016051226,0.00022563213,0.000011086869,0.000060234852,0.00003715937,0.000011264201,0.00032115006],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.97934246,0.013742443,0.0011229677,0.0010772102,0.003676476,0.0010384178],"domain_scores_gemma":[0.8325927,0.12007315,0.020489806,0.003193497,0.020346122,0.0033046342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022348816,0.0004382799,0.00046016445,0.0034023116,0.0025108496,0.005027801,0.00079511746,0.0014228623,0.0010637104],"category_scores_gemma":[0.086984284,0.0003592515,0.00041227235,0.0021789146,0.0019451521,0.0073807137,0.0030200304,0.001466256,0.00030141044],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001906697,0.00037812776,0.5247445,0.0005499778,0.00005088355,0.0007450242,0.3856921,0.00037860844,0.0027796424,0.001590133,0.0017473683,0.08115292],"study_design_scores_gemma":[0.000017002054,0.00066677463,0.32283524,0.00079475966,0.000056570374,0.00051658176,0.64587915,0.00809998,0.0022072867,0.0025285224,0.01628131,0.0001168499],"about_ca_topic_score_codex":0.0042632218,"about_ca_topic_score_gemma":0.0060767964,"teacher_disagreement_score":0.022348816,"about_ca_system_score_codex":0.0029496523,"about_ca_system_score_gemma":0.0024483185,"threshold_uncertainty_score":0.11819327},"labels":[],"label_agreement":null},{"id":"W4391745357","doi":"10.1007/s10664-023-10437-1","title":"A study of common bug fix patterns in Rust","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Rust (programming language); Computer science; Compiler; Programming language; Software engineering","score_opus":0.03218305884190426,"score_gpt":0.3177540607064256,"score_spread":0.28557100186452133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391745357","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9977331,0.00013645248,0.0010290731,0.000055833825,0.0000045118413,0.000018375904,0.00014182861,0.000048460857,0.0008322172],"genre_scores_gemma":[0.9972288,0.00006218439,0.0017525943,0.000024191415,0.000004135053,0.000016335796,0.00026068118,0.0000321706,0.00061873026],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9975048,0.00089597673,0.00019678826,0.00046707335,0.0007669172,0.00016852896],"domain_scores_gemma":[0.9384988,0.037400074,0.013607631,0.004722381,0.004408465,0.0013627138],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.00248267,0.00027496016,0.0003316651,0.0037046296,0.0007818677,0.00082215125,0.00080707046,0.0006708794,0.002519348],"category_scores_gemma":[0.034682855,0.00027829802,0.00045059103,0.0042866752,0.0007667778,0.0014669446,0.00089613075,0.0008551585,0.00034012456],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004437636,0.0008897392,0.8903564,0.00027733468,0.00025490372,0.0014854484,0.012496909,0.002118007,0.005062856,0.0029436252,0.0016110442,0.08205994],"study_design_scores_gemma":[0.00007938253,0.0012986964,0.9694408,0.00011802789,0.00014097827,0.0022801647,0.0069146925,0.011648963,0.002065659,0.0024875272,0.0034747494,0.00005034488],"about_ca_topic_score_codex":0.0039188885,"about_ca_topic_score_gemma":0.008626396,"teacher_disagreement_score":0.99751735,"about_ca_system_score_codex":0.0007726601,"about_ca_system_score_gemma":0.0008518915,"threshold_uncertainty_score":0.013129771},"labels":[],"label_agreement":null},{"id":"W4391835605","doi":"10.1007/s10664-024-10449-5","title":"Quantifying and characterizing clones of self-admitted technical debt in build systems","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Precursory Research for Embryonic Science and Technology; Japan Science and Technology Agency; Japan Society for the Promotion of Science","keywords":"Artifact (error); Technical debt; Computer science; Context (archaeology); Maintainability; Reliability (semiconductor); Code (set theory); Scale (ratio); Data science; Data mining; Software; Software engineering; Software development; Artificial intelligence; Programming language; Biology; Set (abstract data type); Geography","score_opus":0.029443787547574637,"score_gpt":0.29859674966661726,"score_spread":0.2691529621190426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391835605","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9918391,0.00008374062,0.007229241,0.000051053103,0.000003668952,0.00001970081,0.00007459055,0.00009160029,0.0006073612],"genre_scores_gemma":[0.9960174,0.000028632545,0.0035264408,0.000014349905,0.000004980561,0.000016616124,0.000120016186,0.00002932133,0.00024225352],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9951167,0.0012771062,0.00048699882,0.0009049158,0.001739077,0.0004752427],"domain_scores_gemma":[0.87102795,0.06320005,0.032885548,0.016513895,0.013592902,0.0027796132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005711599,0.00041622657,0.0005845566,0.0031042118,0.0009772275,0.0029037127,0.0013624233,0.0015310316,0.0009089044],"category_scores_gemma":[0.08774076,0.0005703214,0.00041925604,0.0024384395,0.0016732505,0.004054311,0.0021511847,0.0012731774,0.00019429578],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001992014,0.0001832686,0.92622787,0.000094139876,0.00011906526,0.00043348392,0.0028852567,0.016506946,0.007964778,0.0084734615,0.000356133,0.036556475],"study_design_scores_gemma":[0.000039304214,0.00038272163,0.7431955,0.00012374127,0.00018627156,0.0012210222,0.003424992,0.20999956,0.014875297,0.024483563,0.0019807103,0.0000872324],"about_ca_topic_score_codex":0.0046688085,"about_ca_topic_score_gemma":0.005823241,"teacher_disagreement_score":0.005711599,"about_ca_system_score_codex":0.0020625128,"about_ca_system_score_gemma":0.0014660542,"threshold_uncertainty_score":0.030206203},"labels":[],"label_agreement":null},{"id":"W4391882875","doi":"10.3390/computers13020052","title":"Interpretable Software Defect Prediction from Project Effort and Static Code Metrics","year":2024,"lang":"en","type":"article","venue":"Computers","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Fanshawe College","funders":"","keywords":"Interpretability; Computer science; Predictive modelling; Machine learning; Software quality; Reliability (semiconductor); Software; Random forest; Software bug; Artificial intelligence; Data mining; Software metric; Support vector machine; Code (set theory); Source code; Quality (philosophy); Reliability engineering; Software development; Set (abstract data type); Engineering; Programming language","score_opus":0.018907956328430994,"score_gpt":0.27081905531226974,"score_spread":0.25191109898383873,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391882875","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8065199,0.00026656775,0.18442518,0.0006493734,0.000027289974,0.00014311926,0.0042537893,0.001027598,0.0026871483],"genre_scores_gemma":[0.9780422,0.00006357295,0.018710477,0.000020783678,0.00000772825,0.000064096006,0.0027343743,0.000032163633,0.00032472162],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99790394,0.0009291078,0.00017528393,0.00039017864,0.0004997276,0.00010184331],"domain_scores_gemma":[0.96582854,0.022671346,0.00510304,0.0032864388,0.002854063,0.00025656528],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043699546,0.00092867634,0.0003811836,0.0036504467,0.00019596989,0.001161087,0.00068324315,0.0007200341,0.0011147753],"category_scores_gemma":[0.032489873,0.00022480186,0.0007025452,0.002037156,0.00037989704,0.0014973177,0.00087765604,0.00095900835,0.00024471624],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033764562,0.000431145,0.53681827,0.0003758954,0.0003182231,0.00041496952,0.001324691,0.25777724,0.0030969584,0.009029426,0.0034213755,0.18665422],"study_design_scores_gemma":[0.000023878618,0.00019799458,0.11937754,0.00008657567,0.0000730315,0.00013263285,0.0004124087,0.8602738,0.002455404,0.0152094215,0.0017148285,0.000042550906],"about_ca_topic_score_codex":0.0034512007,"about_ca_topic_score_gemma":0.0069430657,"teacher_disagreement_score":0.0043699546,"about_ca_system_score_codex":0.0008373971,"about_ca_system_score_gemma":0.00075325806,"threshold_uncertainty_score":0.023110807},"labels":[],"label_agreement":null},{"id":"W4391915512","doi":"10.1007/s10664-024-10443-x","title":"Studying the impact of risk assessment analytics on risk awareness and code review performance","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); University of Waterloo","funders":"","keywords":"Analytics; Computer science; Data science; Code (set theory); Risk analysis (engineering); Business; Programming language","score_opus":0.04283266006473341,"score_gpt":0.3663655825837216,"score_spread":0.3235329225189882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391915512","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99610823,0.00024549014,0.0016103714,0.00031242598,0.000015811345,0.000015418866,0.000053097698,0.00012711773,0.0015119911],"genre_scores_gemma":[0.99877626,0.00004219419,0.00077707536,0.000022359534,0.000011956332,0.000003858007,0.0000540276,0.000015214687,0.00029713794],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99318236,0.0031593712,0.0003971759,0.0008441944,0.0018131196,0.0006037511],"domain_scores_gemma":[0.6318019,0.31279323,0.028002288,0.009355239,0.013681233,0.0043661874],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008920325,0.0005742394,0.00037110742,0.0018758987,0.00042927434,0.0025598418,0.00077103265,0.00092419406,0.0020455488],"category_scores_gemma":[0.14964607,0.00027729318,0.00047421193,0.0014029806,0.00064186525,0.0037879616,0.0007368459,0.0016614302,0.00041534012],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0042182803,0.005520944,0.6534511,0.0003757128,0.0009795391,0.000272735,0.0017011664,0.09593758,0.016034411,0.0043382924,0.0035333969,0.21363679],"study_design_scores_gemma":[0.00016381327,0.0067506833,0.4098689,0.00013157821,0.0006929798,0.00037022756,0.0028930593,0.5476822,0.022240875,0.0071008825,0.0019394966,0.00016521149],"about_ca_topic_score_codex":0.004915294,"about_ca_topic_score_gemma":0.005076402,"teacher_disagreement_score":0.9910797,"about_ca_system_score_codex":0.0012397529,"about_ca_system_score_gemma":0.002392359,"threshold_uncertainty_score":0.047175765},"labels":[],"label_agreement":null},{"id":"W4392347631","doi":"10.1145/3649598","title":"Communicating Study Design Trade-offs in Software Engineering","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Victoria; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Science Foundation Ireland; European Commission; National Science Foundation","keywords":"Computer science; Work (physics); Reflection (computer programming); Process (computing); Risk analysis (engineering); Management science; Strengths and weaknesses; Engineering ethics; Psychology; Business; Engineering; Social psychology","score_opus":0.12728410333757412,"score_gpt":0.3491946088004384,"score_spread":0.22191050546286428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392347631","genre_codex":"methods","genre_gemma":"methods","domain_codex":"methods","domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":"methods","prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06548037,0.01105646,0.7223025,0.14250295,0.006421748,0.02700455,0.0003035513,0.0012804379,0.023647549],"genre_scores_gemma":[0.28138644,0.0018049192,0.6168484,0.030882267,0.0023195036,0.06456007,0.00011165973,0.0005075057,0.0015792415],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.03621928,0.8435737,0.07400297,0.009981275,0.034387335,0.001835453],"domain_scores_gemma":[0.014647949,0.8792941,0.030895647,0.054165337,0.018982721,0.00201419],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.9067478,0.004290369,0.0048906775,0.013317633,0.012499801,0.026581358,0.010011335,0.024767516,0.005899589],"category_scores_gemma":[0.94183356,0.005997211,0.0053601125,0.008446526,0.048024204,0.04351353,0.034710906,0.026730098,0.0024552343],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030426676,0.0006847073,0.013890049,0.011341111,0.0021512657,0.0019907306,0.27467874,0.0046074893,0.006706508,0.27313396,0.017447483,0.39032528],"study_design_scores_gemma":[0.0024257845,0.0026931793,0.0074143745,0.02647537,0.0010199089,0.002178523,0.037261423,0.019695371,0.010564144,0.8026391,0.086474776,0.0011580526],"about_ca_topic_score_codex":0.0009162107,"about_ca_topic_score_gemma":0.0015838072,"teacher_disagreement_score":0.09325218,"about_ca_system_score_codex":0.025189694,"about_ca_system_score_gemma":0.025854979,"threshold_uncertainty_score":0.18276489},"labels":[],"label_agreement":null},{"id":"W4392359402","doi":"10.21203/rs.3.rs-3996923/v1","title":"Interpretation Conclusion Stability of Software Defect Prediction over Time","year":2024,"lang":"en","type":"preprint","venue":"Research Square","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Interpretation (philosophy); Stability (learning theory); Software; Computer science; Machine learning; Programming language","score_opus":0.03091608569368675,"score_gpt":0.3536234347108503,"score_spread":0.3227073490171635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392359402","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7658035,0.0008844391,0.2285926,0.00038235192,0.0001004217,0.00007832628,0.0010341136,0.0020807048,0.0010435352],"genre_scores_gemma":[0.97427994,0.00006529325,0.024407506,0.000029805125,0.000019112043,0.000019800527,0.00085353793,0.000044983288,0.00027993627],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99662125,0.0006323226,0.0003177593,0.0011336359,0.0011142094,0.00018084206],"domain_scores_gemma":[0.97972494,0.010030869,0.002959939,0.0029611222,0.003953023,0.00037000424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043795705,0.00097560696,0.0009242627,0.002607991,0.00033089038,0.0016977135,0.0011309836,0.0008478182,0.00047253474],"category_scores_gemma":[0.027570285,0.0002800144,0.00059970724,0.0015274036,0.00045122823,0.0016528072,0.00073413306,0.0014932133,0.00023960917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092791533,0.00031428522,0.26097104,0.00021387795,0.0003949738,0.0005002359,0.00056538597,0.42906708,0.009553797,0.0011819942,0.0031603673,0.29314902],"study_design_scores_gemma":[0.0000058160886,0.00008698633,0.021335486,0.000018276794,0.00002758506,0.0001087941,0.00008908575,0.9724177,0.0046307747,0.00087815296,0.00038400924,0.000017419228],"about_ca_topic_score_codex":0.007527227,"about_ca_topic_score_gemma":0.0044627986,"teacher_disagreement_score":0.007527227,"about_ca_system_score_codex":0.0010413944,"about_ca_system_score_gemma":0.000772592,"threshold_uncertainty_score":0.02316165},"labels":[],"label_agreement":null},{"id":"W4392518074","doi":"10.1145/3651313","title":"A Psycholinguistics-inspired Method to Counter IP Theft Using Fake Documents","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Management Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"Office of Naval Research","keywords":"Set (abstract data type); Computer science; Deception; Property (philosophy); Psycholinguistics; Artificial intelligence; Natural language processing; Theoretical computer science; Cognition; Programming language; Psychology","score_opus":0.032751265513055974,"score_gpt":0.34429830974899195,"score_spread":0.311547044235936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392518074","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05046487,0.00026652715,0.9428053,0.00083166885,0.000057524234,0.00017545182,0.00008756329,0.0016352544,0.0036758326],"genre_scores_gemma":[0.623235,0.0001399577,0.37314627,0.00025606414,0.000049576436,0.00013767778,0.00014077882,0.00014204164,0.0027526433],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99834955,0.0006405686,0.000096399424,0.0003058558,0.0004832223,0.00012446193],"domain_scores_gemma":[0.9943197,0.0032559712,0.0007231921,0.0009967476,0.0005582999,0.00014612665],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002296894,0.0010578799,0.0009413191,0.0013255456,0.000735739,0.0017027031,0.0014993193,0.0015372875,0.0025124194],"category_scores_gemma":[0.011218899,0.00047547751,0.0008709556,0.0005342998,0.0017831686,0.002585718,0.0013503175,0.0019039623,0.00065761665],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047441723,0.0006593198,0.0038096344,0.00045156627,0.0002326107,0.0003754549,0.0004999806,0.44155762,0.047081094,0.06517656,0.0048777713,0.43480393],"study_design_scores_gemma":[0.000028121329,0.0001347366,0.0005075502,0.000017042628,0.000027673628,0.00020610502,0.000055760473,0.96102756,0.011275975,0.024960987,0.001725614,0.000032916505],"about_ca_topic_score_codex":0.0011477191,"about_ca_topic_score_gemma":0.0020572345,"teacher_disagreement_score":0.0025124194,"about_ca_system_score_codex":0.0011791729,"about_ca_system_score_gemma":0.001542473,"threshold_uncertainty_score":0.012147307},"labels":[],"label_agreement":null},{"id":"W4392591292","doi":"10.1007/s11219-024-09663-7","title":"A comprehensive catalog of refactoring strategies to handle test smells in Java-based systems","year":2024,"lang":"en","type":"article","venue":"Software Quality Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Instituto Nacional de Ciência e Tecnologia para Engenharia de Software; Fundação de Amparo à Pesquisa do Estado da Bahia; Conselho Nacional de Desenvolvimento Científico e Tecnológico; Coordenação de Aperfeiçoamento de Pessoal de Nível Superior","keywords":"Code refactoring; Java; Computer science; Programming language; Software engineering; Operating system; Software","score_opus":0.05475254564650548,"score_gpt":0.3527974695137458,"score_spread":0.2980449238672403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392591292","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058454152,0.26150396,0.51963115,0.0025399558,0.00082227733,0.0022949532,0.030978244,0.05681363,0.066961735],"genre_scores_gemma":[0.07562591,0.17632598,0.6590689,0.001532899,0.0003528671,0.0008959882,0.05712475,0.0062906714,0.022782002],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9967457,0.00036934743,0.0008684666,0.00028781008,0.0015472542,0.00018133581],"domain_scores_gemma":[0.9871975,0.006871643,0.0011493817,0.0020897354,0.002410618,0.00028109015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003068578,0.002052074,0.0015500144,0.014133684,0.0007599841,0.0021834613,0.0024836278,0.0014744595,0.008487599],"category_scores_gemma":[0.012936805,0.0012478057,0.0020948946,0.011089885,0.00031090973,0.003157371,0.0011475611,0.00097265927,0.004832172],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007899675,0.000220956,0.0013385419,0.0044590817,0.00011916102,0.00013012618,0.00013404075,0.0017336083,0.008317524,0.0012304853,0.015640195,0.9665973],"study_design_scores_gemma":[0.0003475355,0.0014024233,0.029491214,0.018802388,0.002452952,0.004868132,0.00048794082,0.045278095,0.11491415,0.016011128,0.76524425,0.0006998083],"about_ca_topic_score_codex":0.0026762711,"about_ca_topic_score_gemma":0.006799267,"teacher_disagreement_score":0.014133684,"about_ca_system_score_codex":0.0005071242,"about_ca_system_score_gemma":0.002319932,"threshold_uncertainty_score":0.028393865},"labels":[],"label_agreement":null},{"id":"W4392619070","doi":"10.1145/3643991.3644926","title":"Whodunit: Classifying Code as Human Authored or GPT-4 Generated - A case study on CodeChef problems","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Classifier (UML); Computer science; Machine learning; Artificial intelligence; Genetic programming; Stylometry; Natural language processing; Code (set theory); Source code; Receiver operating characteristic; Programming language","score_opus":0.1743010673399469,"score_gpt":0.40177121112458075,"score_spread":0.22747014378463384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392619070","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9224656,0.0020490154,0.033798132,0.0017516899,0.0004676016,0.00051107444,0.013834654,0.010783212,0.014339091],"genre_scores_gemma":[0.8743815,0.00056774094,0.07536564,0.00062461925,0.00014747975,0.0003688863,0.037237667,0.0015135087,0.009793051],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9960056,0.0009884038,0.0003739695,0.00091716857,0.0014542084,0.00026060783],"domain_scores_gemma":[0.9776585,0.013375507,0.001613571,0.0035249782,0.00294722,0.00088019716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027337994,0.0008728382,0.00048504208,0.004089516,0.000982833,0.0017626724,0.0012180546,0.0021192066,0.001994217],"category_scores_gemma":[0.026642112,0.00023466101,0.0006686692,0.002499035,0.001266381,0.0025384354,0.0014191902,0.0017818571,0.0018103963],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016499937,0.0015084767,0.19855201,0.0021208492,0.00019215018,0.003351633,0.0045868577,0.035322998,0.013489639,0.010634418,0.16907765,0.55951333],"study_design_scores_gemma":[0.00032294937,0.0009945531,0.13082328,0.0006205547,0.00009709062,0.00638376,0.003956012,0.6023185,0.04754247,0.020818587,0.18588255,0.00023963244],"about_ca_topic_score_codex":0.005247887,"about_ca_topic_score_gemma":0.010756041,"teacher_disagreement_score":0.005247887,"about_ca_system_score_codex":0.0016068729,"about_ca_system_score_gemma":0.0011122486,"threshold_uncertainty_score":0.014457881},"labels":[],"label_agreement":null},{"id":"W4392907063","doi":"10.32920/25412866.v1","title":"RLML: A Domain-specific Modelling Language for Reinforcement Learning","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Reinforcement learning; Syntax; Domain (mathematical analysis); Abstraction; Artificial intelligence; Domain-specific language; Simplicity; Modeling language; Constraint (computer-aided design); Software engineering; Machine learning; Popularity; Programming language; Human–computer interaction; Software; Engineering","score_opus":0.035770806054210603,"score_gpt":0.28991364152208215,"score_spread":0.25414283546787153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392907063","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004656957,0.00007369592,0.98320526,0.0002701483,0.000074004536,0.00013749013,0.0008834924,0.01319549,0.0016947191],"genre_scores_gemma":[0.029745402,0.00043582256,0.9527943,0.0005885435,0.00009180863,0.0012375799,0.0035891386,0.0061151874,0.0054022055],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970739,0.0010616513,0.00054118736,0.000398296,0.00074614113,0.0001788423],"domain_scores_gemma":[0.9945832,0.003209508,0.00041884807,0.0008049963,0.00082214887,0.0001613464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004133691,0.0017247661,0.00095521845,0.0010211043,0.00070389535,0.003594622,0.0042313463,0.0029701162,0.014706438],"category_scores_gemma":[0.010941006,0.0017236426,0.0025058852,0.0007696002,0.0017551982,0.0031992337,0.0025248278,0.005513063,0.006492706],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037753023,0.00026004313,0.0012697977,0.0022186653,0.00015901387,0.0010296098,0.0018177802,0.19872776,0.01622469,0.52579135,0.08124296,0.17088078],"study_design_scores_gemma":[0.0001864973,0.0001226095,0.00017407506,0.00045200638,0.000060878836,0.0006098637,0.00010450048,0.4172414,0.014425441,0.112054326,0.45443085,0.00013759393],"about_ca_topic_score_codex":0.0041610403,"about_ca_topic_score_gemma":0.004807334,"teacher_disagreement_score":0.014706438,"about_ca_system_score_codex":0.0012168474,"about_ca_system_score_gemma":0.003352384,"threshold_uncertainty_score":0.049197912},"labels":[],"label_agreement":null},{"id":"W4392907203","doi":"10.32920/25412866","title":"RLML: A Domain-specific Modelling Language for Reinforcement Learning","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Reinforcement learning; Syntax; Domain (mathematical analysis); Abstraction; Artificial intelligence; Domain-specific language; Simplicity; Constraint (computer-aided design); Modeling language; Machine learning; Programming language; Software engineering; Human–computer interaction; Software; Engineering","score_opus":0.035770806054210603,"score_gpt":0.28991364152208215,"score_spread":0.25414283546787153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392907203","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004656957,0.00007369592,0.98320526,0.0002701483,0.000074004536,0.00013749013,0.0008834924,0.01319549,0.0016947191],"genre_scores_gemma":[0.029745402,0.00043582256,0.9527943,0.0005885435,0.00009180863,0.0012375799,0.0035891386,0.0061151874,0.0054022055],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970739,0.0010616513,0.00054118736,0.000398296,0.00074614113,0.0001788423],"domain_scores_gemma":[0.9945832,0.003209508,0.00041884807,0.0008049963,0.00082214887,0.0001613464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004133691,0.0017247661,0.00095521845,0.0010211043,0.00070389535,0.003594622,0.0042313463,0.0029701162,0.014706438],"category_scores_gemma":[0.010941006,0.0017236426,0.0025058852,0.0007696002,0.0017551982,0.0031992337,0.0025248278,0.005513063,0.006492706],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037753023,0.00026004313,0.0012697977,0.0022186653,0.00015901387,0.0010296098,0.0018177802,0.19872776,0.01622469,0.52579135,0.08124296,0.17088078],"study_design_scores_gemma":[0.0001864973,0.0001226095,0.00017407506,0.00045200638,0.000060878836,0.0006098637,0.00010450048,0.4172414,0.014425441,0.112054326,0.45443085,0.00013759393],"about_ca_topic_score_codex":0.0041610403,"about_ca_topic_score_gemma":0.004807334,"teacher_disagreement_score":0.014706438,"about_ca_system_score_codex":0.0012168474,"about_ca_system_score_gemma":0.003352384,"threshold_uncertainty_score":0.049197912},"labels":[],"label_agreement":null},{"id":"W4393035092","doi":"10.1109/iwsc60764.2023.00012","title":"TransClone: A Language Agnostic Code Clone Detector","year":2023,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Programming language; Code (set theory); clone (Java method); Detector; Natural language processing; Telecommunications; Set (abstract data type); Biology","score_opus":0.01803039058598823,"score_gpt":0.27923751451598666,"score_spread":0.26120712392999845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393035092","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08229727,0.0031928855,0.50175035,0.0007763682,0.0006338573,0.0007072432,0.029728351,0.3719537,0.008959931],"genre_scores_gemma":[0.21646897,0.0011046397,0.63880855,0.001454397,0.00017577849,0.0009790803,0.11065168,0.012950054,0.017406762],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974974,0.00022843522,0.00017284224,0.0009812035,0.0009484192,0.00017163674],"domain_scores_gemma":[0.99627185,0.0011368067,0.0003464275,0.0010096581,0.00110029,0.00013501324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015338316,0.0024727266,0.0012287339,0.005809843,0.0006270273,0.0019387142,0.0035165024,0.0019565392,0.0026571506],"category_scores_gemma":[0.00837499,0.000786485,0.0017561427,0.0027096674,0.0007234921,0.0039083655,0.0030351197,0.0017418177,0.0038063722],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006512128,0.00045010954,0.033652835,0.0015606863,0.0005256613,0.0010620552,0.0005726043,0.02458147,0.046981435,0.006876575,0.22608563,0.65699965],"study_design_scores_gemma":[0.00021853035,0.00045966855,0.013541362,0.0002564658,0.00030415546,0.0030202894,0.00033783307,0.63605917,0.1339668,0.02011007,0.19149223,0.00023338964],"about_ca_topic_score_codex":0.0073263184,"about_ca_topic_score_gemma":0.011350172,"teacher_disagreement_score":0.0073263184,"about_ca_system_score_codex":0.0012257624,"about_ca_system_score_gemma":0.0015156564,"threshold_uncertainty_score":0.014567316},"labels":[],"label_agreement":null},{"id":"W4393142189","doi":"10.1007/s00766-024-00416-3","title":"Improving requirements completeness: automated assistance through large language models","year":2024,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Terminology; Completeness (order theory); Computer science; Natural language; Artificial intelligence; Filter (signal processing); Process (computing); Noise (video); Software engineering; Natural language processing; Programming language; Linguistics","score_opus":0.03724158375314305,"score_gpt":0.3098539232716619,"score_spread":0.2726123395185188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393142189","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023645474,0.00013350103,0.9614684,0.00080254575,0.000035943478,0.0002823353,0.0008285944,0.010568277,0.0022349786],"genre_scores_gemma":[0.28422284,0.00023480848,0.7078695,0.0003230471,0.00004014763,0.00035424772,0.0031853106,0.0019510679,0.0018189745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9883581,0.005529559,0.0006369453,0.0008645871,0.004248706,0.00036214618],"domain_scores_gemma":[0.95073277,0.03325485,0.002246882,0.008317512,0.005059934,0.00038812775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064913654,0.0013859167,0.0012394286,0.001959548,0.0010732911,0.003309491,0.0029069162,0.0017075398,0.0050056675],"category_scores_gemma":[0.048905995,0.0016126159,0.0023567518,0.0013580907,0.000984143,0.0069764266,0.004180741,0.003104442,0.002150938],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008593555,0.0013686217,0.0076597426,0.0015890006,0.00033766567,0.0012067262,0.002340688,0.33222428,0.04051653,0.06918021,0.024621308,0.5180959],"study_design_scores_gemma":[0.00008650822,0.00007194664,0.0003194573,0.00009109673,0.00009343418,0.00015730066,0.00023409558,0.95118964,0.012396377,0.026818044,0.008500191,0.00004190691],"about_ca_topic_score_codex":0.005373096,"about_ca_topic_score_gemma":0.014823735,"teacher_disagreement_score":0.0064913654,"about_ca_system_score_codex":0.0013546706,"about_ca_system_score_gemma":0.0048671816,"threshold_uncertainty_score":0.03433001},"labels":[],"label_agreement":null},{"id":"W4393155044","doi":"10.1007/978-3-031-49179-5_24","title":"How Is Software Reuse Discussed in Stack Overflow?","year":2024,"lang":"en","type":"book-chapter","venue":"Conference on systems engineering research series","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Reuse; Computer science; Stack (abstract data type); Operating system; Engineering; Waste management","score_opus":0.07593201115550784,"score_gpt":0.3012310783105806,"score_spread":0.22529906715507275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393155044","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18116242,0.035016228,0.16512911,0.09476411,0.003052361,0.000056741184,0.00024723078,0.0013471213,0.5192247],"genre_scores_gemma":[0.87270534,0.013292852,0.020004617,0.0041248095,0.0020604378,0.000045765897,0.00017928996,0.00092528824,0.08666166],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9974126,0.0006989897,0.00012238187,0.0002372604,0.0010116842,0.0005170904],"domain_scores_gemma":[0.99416625,0.0033185312,0.00065532536,0.00073063053,0.00081178895,0.00031740643],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.002998457,0.0005697927,0.0005510983,0.0033250814,0.0031997887,0.009153585,0.001338117,0.0037345276,0.007455229],"category_scores_gemma":[0.016810084,0.0005590791,0.00085981895,0.004944353,0.0068597393,0.029557448,0.0050414996,0.0034970285,0.0010736658],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000613608,0.000036852012,0.0023294087,0.0001881218,0.000035659516,0.00059375144,0.009141299,0.0009539256,0.001029481,0.8295865,0.017399123,0.13864441],"study_design_scores_gemma":[0.000011190017,0.000035301717,0.0024615764,0.0006515009,0.00007945698,0.0012840405,0.0075843437,0.0036091,0.0024044153,0.806181,0.17564225,0.000055851247],"about_ca_topic_score_codex":0.0070891683,"about_ca_topic_score_gemma":0.0066177533,"teacher_disagreement_score":0.9970015,"about_ca_system_score_codex":0.0028890413,"about_ca_system_score_gemma":0.0030894612,"threshold_uncertainty_score":0.024940252},"labels":[],"label_agreement":null},{"id":"W4393213239","doi":"10.1145/3597503.3639194","title":"ChatGPT Incorrectness Detection in Software Reviews","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Suite; Benchmark (surveying); Selection (genetic algorithm); Generative grammar; Artificial intelligence; Software; Machine learning; Natural language processing; Information retrieval; Programming language","score_opus":0.02675092483946025,"score_gpt":0.29921989749225403,"score_spread":0.2724689726527938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393213239","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9606348,0.0017536292,0.029994495,0.0004583587,0.0001287622,0.0004626432,0.0011985023,0.0025032375,0.0028656097],"genre_scores_gemma":[0.9705085,0.00039964446,0.02489638,0.00022887424,0.00006553668,0.00028197287,0.001953646,0.00019691282,0.001468543],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.95090973,0.022082211,0.004356717,0.00537654,0.016111717,0.0011630866],"domain_scores_gemma":[0.6284948,0.27506426,0.040485434,0.010735309,0.042914953,0.0023052576],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.017003447,0.0009707248,0.0011753791,0.008435271,0.0009967615,0.0015698846,0.0013851303,0.0013889604,0.0009627976],"category_scores_gemma":[0.20541456,0.00052257936,0.00048258147,0.004014659,0.0005735773,0.0025145093,0.0018833517,0.0011430891,0.00086007704],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010508727,0.0005272672,0.54754996,0.0036633476,0.00034625133,0.0021942523,0.02057349,0.0047988067,0.028452884,0.0013999034,0.013056967,0.37638602],"study_design_scores_gemma":[0.00014745444,0.0017565164,0.73369056,0.0013620721,0.00045892116,0.0061168578,0.010200632,0.14825213,0.05659498,0.0039040623,0.03711164,0.0004041562],"about_ca_topic_score_codex":0.002053674,"about_ca_topic_score_gemma":0.003936095,"teacher_disagreement_score":0.9829966,"about_ca_system_score_codex":0.0009971063,"about_ca_system_score_gemma":0.0009908275,"threshold_uncertainty_score":0.08992392},"labels":[],"label_agreement":null},{"id":"W4393284497","doi":"10.1145/3650105.3652301","title":"Exploring the Impact of the Output Format on the Evaluation of Large Language Models for Code Translation","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Queen's University","funders":"Vector Institute","keywords":"Computer science; Python (programming language); Programming language; Java; Source code; Software engineering; KPI-driven code analysis; Code (set theory); Software; Natural language processing; Artificial intelligence; Software development; Static program analysis","score_opus":0.2895498212732952,"score_gpt":0.4019542227165666,"score_spread":0.1124044014432714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393284497","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7997981,0.003032499,0.1488027,0.0013184972,0.00067554845,0.0006568291,0.0048339283,0.02830791,0.012574064],"genre_scores_gemma":[0.85526127,0.0006634594,0.124032415,0.0005315177,0.000069280504,0.00045869662,0.011420106,0.0051098354,0.002453329],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9837337,0.009004187,0.0015073086,0.0019418009,0.0032988994,0.0005141765],"domain_scores_gemma":[0.9352323,0.046873733,0.0021638789,0.008812395,0.0062266155,0.0006910934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0152581185,0.0023825048,0.0010502177,0.0019605872,0.0008802413,0.0031802226,0.0018872145,0.0014947256,0.0024067573],"category_scores_gemma":[0.09565778,0.00066593045,0.0012849921,0.0018126069,0.0011478323,0.0041937483,0.0025657413,0.002477839,0.0022492937],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004790709,0.0016331527,0.07591599,0.0040375767,0.0010653393,0.0011041132,0.0036393937,0.28354493,0.049673874,0.009221428,0.028547138,0.5368263],"study_design_scores_gemma":[0.00037342115,0.0014834437,0.013428765,0.0005035189,0.00041346252,0.0005836737,0.0013780324,0.877745,0.07541198,0.007863635,0.020638576,0.00017648705],"about_ca_topic_score_codex":0.005201997,"about_ca_topic_score_gemma":0.00629015,"teacher_disagreement_score":0.0152581185,"about_ca_system_score_codex":0.001538519,"about_ca_system_score_gemma":0.0020349757,"threshold_uncertainty_score":0.0806936},"labels":[],"label_agreement":null},{"id":"W4393334908","doi":"10.1145/3643916.3644413","title":"Rationale Dataset and Analysis for the Commit Messages of the Linux Kernel Out-of-Memory Killer","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal","funders":"Fonds de recherche du Québec – Nature et technologies","keywords":"Commit; Computer science; Kernel (algebra); Source code; Component (thermodynamics); Code (set theory); Linux kernel; Programming language; Operating system; Database; Set (abstract data type)","score_opus":0.047856793625842595,"score_gpt":0.31735882014186395,"score_spread":0.26950202651602134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393334908","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22973877,0.0019116565,0.009276998,0.002471323,0.0007168182,0.00080011925,0.73653436,0.009040202,0.009509789],"genre_scores_gemma":[0.055160936,0.00023652322,0.015447857,0.00041600407,0.00012667362,0.0006323723,0.9249206,0.0004375795,0.0026214514],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99666035,0.00075393677,0.00046315443,0.0007231373,0.0011413418,0.0002581917],"domain_scores_gemma":[0.9872564,0.006232274,0.0014603375,0.0016429605,0.0026274049,0.0007805912],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024258383,0.00113937,0.00056158926,0.005732268,0.0018948212,0.0012067511,0.001309224,0.003193814,0.0031708283],"category_scores_gemma":[0.0133490795,0.00041362902,0.00091920415,0.0031187923,0.0009053684,0.0014973946,0.0017169734,0.0023894645,0.0044934754],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012710241,0.0014912717,0.041592203,0.0046264525,0.00020463542,0.0024943324,0.004025086,0.0038855968,0.031037962,0.004699898,0.80472755,0.099944025],"study_design_scores_gemma":[0.00073745527,0.0006997645,0.24396928,0.00088097143,0.00019579077,0.0028251414,0.0045742155,0.03333208,0.021452803,0.005485814,0.68548363,0.00036309482],"about_ca_topic_score_codex":0.013746794,"about_ca_topic_score_gemma":0.033770345,"teacher_disagreement_score":0.013746794,"about_ca_system_score_codex":0.0013230263,"about_ca_system_score_gemma":0.0018589081,"threshold_uncertainty_score":0.027333558},"labels":[],"label_agreement":null},{"id":"W4393372268","doi":"10.1109/ms.2024.3382364","title":"Toward Optimal Psychological Functioning in AI-Driven Software Engineering Tasks: The Software Evaluation for Well-Being and Optimal Psychological Functioning in a Context-Aware Environment Assessment Framework","year":2024,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Context (archaeology); Software engineering; Software development; Software; Psychological testing; Psychology; Clinical psychology; Programming language","score_opus":0.0358834102836699,"score_gpt":0.3324900582465691,"score_spread":0.29660664796289915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393372268","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7839263,0.0011613317,0.17671631,0.0038743117,0.00009235974,0.0006926241,0.00020802989,0.00023068638,0.03309811],"genre_scores_gemma":[0.9612083,0.0001769992,0.037574418,0.00016922288,0.000010173204,0.00030921685,0.000053556258,0.000017445636,0.00048064097],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9930767,0.004270391,0.0003042807,0.00046103046,0.0015382735,0.0003493318],"domain_scores_gemma":[0.9907354,0.0037025558,0.0016627229,0.0007223562,0.0019740195,0.001203045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008432878,0.0007055525,0.0003882974,0.0026456234,0.0010869929,0.0031978912,0.0005785716,0.0009225277,0.0008927036],"category_scores_gemma":[0.01875875,0.0002081072,0.00064545084,0.0011554252,0.0030566305,0.0027126907,0.0036007909,0.0015944557,0.00016417763],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006189021,0.0025044892,0.4033606,0.00085267285,0.00035917308,0.00038196275,0.03604646,0.011631112,0.011100996,0.15209281,0.0052894633,0.37576136],"study_design_scores_gemma":[0.00009578993,0.002883191,0.7046961,0.0009414747,0.00027105806,0.00063161855,0.031957757,0.06656476,0.009247035,0.16358802,0.018764833,0.00035836527],"about_ca_topic_score_codex":0.0019302369,"about_ca_topic_score_gemma":0.0031108889,"teacher_disagreement_score":0.008432878,"about_ca_system_score_codex":0.0019129735,"about_ca_system_score_gemma":0.0021531938,"threshold_uncertainty_score":0.044597864},"labels":[],"label_agreement":null},{"id":"W4393416783","doi":"10.5281/zenodo.7549218","title":"Bugsplainer: Explaining Software Bugs Leveraging Code Structures in Neural Machine Translation","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Code (set theory); Programming language; Translation (biology); Machine translation; Software bug; Software; Software engineering; Operating system; Artificial intelligence; Biology","score_opus":0.04739010808059809,"score_gpt":0.26718827025099434,"score_spread":0.21979816217039624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393416783","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1693719,0.008616614,0.08887042,0.0039379653,0.0008492367,0.0010515711,0.61871004,0.095821,0.012771297],"genre_scores_gemma":[0.09175816,0.0009233145,0.07584169,0.0005565897,0.00007377857,0.00054669793,0.8254523,0.00085903006,0.0039884658],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99893445,0.000311351,0.00009348122,0.00037152303,0.00020740427,0.00008179379],"domain_scores_gemma":[0.9971776,0.0015738528,0.00022613477,0.0005597626,0.00038002103,0.00008260444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001290453,0.0026408175,0.00058421434,0.0040319264,0.00082811073,0.0009520031,0.0025376568,0.0029114624,0.0057114563],"category_scores_gemma":[0.0078048846,0.00054294,0.0016886869,0.0024377923,0.0005995636,0.0018640662,0.0017665505,0.0020572324,0.005269921],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008883722,0.0007952684,0.026861744,0.003501658,0.00037937044,0.0018917955,0.0006911488,0.035745554,0.0062118084,0.0060257153,0.6947535,0.22225405],"study_design_scores_gemma":[0.0015302217,0.0006115472,0.029330472,0.0006528098,0.00041365225,0.002289394,0.00074480916,0.5456747,0.01912954,0.034479473,0.3648967,0.00024671183],"about_ca_topic_score_codex":0.023586364,"about_ca_topic_score_gemma":0.06813236,"teacher_disagreement_score":0.023586364,"about_ca_system_score_codex":0.0015925461,"about_ca_system_score_gemma":0.0018146063,"threshold_uncertainty_score":0.046898127},"labels":[],"label_agreement":null},{"id":"W4393422503","doi":"10.5281/zenodo.5329118","title":"Breakpoints into the wild: an exploratory study","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Breakpoint; Biology; Genetics; Gene","score_opus":0.03377372250589496,"score_gpt":0.266218491671774,"score_spread":0.23244476916587908,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393422503","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006800652,0.0006298354,0.00047384133,0.00019233208,0.00004790411,0.000069224756,0.98939157,0.00048760773,0.001907049],"genre_scores_gemma":[0.004549137,0.00023483069,0.00111485,0.00008592744,0.0000129392965,0.00022055273,0.99267113,0.000111156056,0.000999499],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972494,0.00057542656,0.00048490494,0.00072933116,0.00065950525,0.0003013279],"domain_scores_gemma":[0.9912719,0.003791967,0.0009710074,0.0017133562,0.0017463294,0.0005054943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021275026,0.001046486,0.0008956084,0.008103475,0.0010290282,0.0021090796,0.0018427015,0.0017242162,0.014731797],"category_scores_gemma":[0.011132066,0.00037245604,0.0010314898,0.011495837,0.0005006662,0.0017231765,0.0024374642,0.0013508603,0.016917618],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058506354,0.00022339924,0.023555323,0.0037905811,0.0001261645,0.00044711886,0.0006558966,0.0007106699,0.00093887455,0.0021503333,0.9423252,0.024491454],"study_design_scores_gemma":[0.00024394272,0.00008379503,0.037561618,0.0009272642,0.00007704556,0.00044938113,0.0011958575,0.000934626,0.0011017345,0.0014526919,0.9559114,0.00006069192],"about_ca_topic_score_codex":0.01142823,"about_ca_topic_score_gemma":0.029017117,"teacher_disagreement_score":0.014731797,"about_ca_system_score_codex":0.0014351323,"about_ca_system_score_gemma":0.0016074888,"threshold_uncertainty_score":0.04928279},"labels":[],"label_agreement":null},{"id":"W4393438819","doi":"10.5281/zenodo.7812452","title":"Dataset for \"Deep Dive into the Verifiability of Code Generated by GitHub Copilot\"","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code (set theory); Computer science; Computer security; Programming language","score_opus":0.03581037201718835,"score_gpt":0.2719287974016264,"score_spread":0.23611842538443806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393438819","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011352075,0.00022894816,0.00041454585,0.0001642958,0.00007286842,0.000040092367,0.99182093,0.0043271463,0.0017960444],"genre_scores_gemma":[0.0010131922,0.000047536323,0.0005606423,0.0000744467,0.0000067076576,0.00007108836,0.99762946,0.00018612163,0.00041085508],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979746,0.00031855426,0.0002292837,0.0004827912,0.0007400211,0.00025488058],"domain_scores_gemma":[0.99628323,0.0012866624,0.00029490024,0.0011118692,0.00076217175,0.0002611859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012905255,0.004238026,0.001445764,0.0033783077,0.0012495558,0.0019408885,0.0041390755,0.0039679916,0.028012957],"category_scores_gemma":[0.0076814955,0.00075843406,0.0017876669,0.0033609984,0.0008768523,0.0013777339,0.002682678,0.002652802,0.049461327],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015396408,0.00006246536,0.0007981038,0.0010724752,0.000048961683,0.0001024657,0.00003182611,0.0009607631,0.0004560722,0.0007626752,0.9917018,0.0038485446],"study_design_scores_gemma":[0.0010706183,0.00010762543,0.0074791084,0.00061009236,0.000093481,0.0003898773,0.00012060778,0.0062642586,0.0030623206,0.0050821495,0.97562814,0.00009174235],"about_ca_topic_score_codex":0.022203594,"about_ca_topic_score_gemma":0.046191867,"teacher_disagreement_score":0.028012957,"about_ca_system_score_codex":0.0018999478,"about_ca_system_score_gemma":0.0027607044,"threshold_uncertainty_score":0.09371269},"labels":[],"label_agreement":null},{"id":"W4393467273","doi":"10.5281/zenodo.4052320","title":"Empirical analysis of Type-Related Defects in Python projects","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; McGill University","funders":"","keywords":"Python (programming language); Programming language; Computer science","score_opus":0.04724435252124957,"score_gpt":0.285587535010376,"score_spread":0.23834318248912642,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393467273","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037924528,0.0003180716,0.0004186766,0.0002916224,0.00007237087,0.00005552202,0.958978,0.00043497173,0.0015062457],"genre_scores_gemma":[0.016403995,0.00008313033,0.00051944394,0.00005201444,0.000015230548,0.00012083585,0.9818127,0.00004707954,0.0009456453],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99734014,0.00055654213,0.0003116026,0.0006728081,0.00076663995,0.00035228176],"domain_scores_gemma":[0.9904456,0.0026571613,0.0018251198,0.002247471,0.0021172962,0.0007072419],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00277154,0.0009424871,0.00078317866,0.004167685,0.00070400396,0.0012548915,0.0018516637,0.001337466,0.007884253],"category_scores_gemma":[0.012359186,0.0003770052,0.0009236539,0.004844403,0.00046417522,0.0011039444,0.0022797654,0.0014676878,0.012021337],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083336455,0.00055643514,0.16044039,0.001594192,0.0002687944,0.00026178343,0.0003047853,0.0021099513,0.00083827827,0.0013755908,0.8076624,0.023753937],"study_design_scores_gemma":[0.00057892286,0.0002684445,0.6679495,0.0005721444,0.00017726423,0.0010089327,0.0012511868,0.0062164147,0.0022383905,0.002036752,0.31754953,0.00015259403],"about_ca_topic_score_codex":0.013126014,"about_ca_topic_score_gemma":0.022096539,"teacher_disagreement_score":0.013126014,"about_ca_system_score_codex":0.0010547568,"about_ca_system_score_gemma":0.0013235013,"threshold_uncertainty_score":0.026375473},"labels":[],"label_agreement":null},{"id":"W4393512234","doi":"10.5281/zenodo.4719160","title":"Automated Support for Searching and Selecting Evidence in Software Engineering: A Cross-domain Systematic Mapping","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Computer science; Software engineering; Domain (mathematical analysis); Software; Systems engineering; Data mining; Engineering; Programming language; Mathematics","score_opus":0.043030932149270076,"score_gpt":0.28234802919294694,"score_spread":0.23931709704367687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393512234","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015224752,0.0058244397,0.0011883867,0.00030628676,0.00003072254,0.0014978446,0.98812383,0.0002884265,0.0012175428],"genre_scores_gemma":[0.0145165445,0.009199218,0.023230812,0.0008284982,0.000040425904,0.024709854,0.92569906,0.00030379544,0.0014717035],"study_design_codex":"systematic_review","study_design_gemma":"not_applicable","domain_scores_codex":[0.979791,0.005906399,0.009785433,0.002017717,0.002039463,0.00045998048],"domain_scores_gemma":[0.8842937,0.0831794,0.012920547,0.0075270645,0.01081364,0.0012655967],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013090083,0.0017311926,0.0047300723,0.023319244,0.0013678058,0.0034494707,0.0027537802,0.0026166865,0.05119953],"category_scores_gemma":[0.116087385,0.0014292016,0.0051020035,0.020891663,0.00065929786,0.0021943185,0.005578764,0.0018415072,0.008579708],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002604446,0.00016267979,0.009231255,0.56711054,0.009385158,0.00039982164,0.0011735375,0.0017657774,0.002000579,0.0036359609,0.32985762,0.07267269],"study_design_scores_gemma":[0.008688099,0.00039503086,0.03796126,0.1759106,0.021361193,0.00062899024,0.0008846657,0.0010722069,0.0023050827,0.0076719485,0.74280274,0.00031819212],"about_ca_topic_score_codex":0.011214446,"about_ca_topic_score_gemma":0.038473617,"teacher_disagreement_score":0.9869099,"about_ca_system_score_codex":0.0034305626,"about_ca_system_score_gemma":0.011770925,"threshold_uncertainty_score":0.17127949},"labels":[],"label_agreement":null},{"id":"W4393525195","doi":"10.5281/zenodo.3633080","title":"Empirical Study of the Relationship between Design Patterns and Code Smells","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code smell; Code (set theory); Empirical research; Computer science; Programming language; Mathematics; Statistics; Software quality; Software","score_opus":0.13376646941934978,"score_gpt":0.3097928022014841,"score_spread":0.1760263327821343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393525195","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99756694,0.0000880562,0.0015828192,0.000045574296,0.0000036558163,0.000046561716,0.00023763839,0.000018816223,0.0004099898],"genre_scores_gemma":[0.9973642,0.000036213292,0.001932411,0.000012680267,0.0000053694844,0.0000774254,0.0004113448,0.000010744087,0.0001495534],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98392165,0.006164618,0.0024838801,0.0024466505,0.004505711,0.00047750684],"domain_scores_gemma":[0.5153909,0.3780537,0.06246068,0.019530034,0.021944126,0.0026205776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013941171,0.000426702,0.00038211377,0.0032272006,0.0005287691,0.0011646765,0.0007461306,0.0007803991,0.0011030644],"category_scores_gemma":[0.15973903,0.00039194073,0.00060123595,0.0033123174,0.0012867956,0.0020474282,0.0014741396,0.0011697196,0.0002469968],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016617906,0.00024037855,0.9833182,0.00016108052,0.00015299307,0.000121125064,0.0021035175,0.0006285213,0.0008531946,0.00017639415,0.00018739027,0.011890986],"study_design_scores_gemma":[0.000024977699,0.00050851854,0.98745406,0.00005193799,0.00007573288,0.00036098552,0.0024572418,0.005980544,0.0015672835,0.0004165444,0.0010737358,0.000028329201],"about_ca_topic_score_codex":0.0011436979,"about_ca_topic_score_gemma":0.0015897062,"teacher_disagreement_score":0.013941171,"about_ca_system_score_codex":0.00064455485,"about_ca_system_score_gemma":0.00061497494,"threshold_uncertainty_score":0.0737288},"labels":[],"label_agreement":null},{"id":"W4393529708","doi":"10.5281/zenodo.8190493","title":"Integrating Visual Aids to Enhance the Code Reviewer Selection Process Replication Package","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Replication (statistics); Computer science; Code (set theory); Selection (genetic algorithm); Process (computing); R package; Psychology; Programming language; Biology; Artificial intelligence; Virology","score_opus":0.02873166098799828,"score_gpt":0.3243248768035528,"score_spread":0.29559321581555453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393529708","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0750743,0.002285375,0.03719515,0.0017951052,0.00067488523,0.0026518703,0.8365082,0.034049146,0.009765935],"genre_scores_gemma":[0.057513013,0.00053192436,0.09328509,0.00040733954,0.00012740826,0.0035719653,0.8397211,0.00080859597,0.004033589],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99411523,0.0023102087,0.0010587957,0.001177652,0.0010382434,0.00029987484],"domain_scores_gemma":[0.9693554,0.012689481,0.0029916088,0.006462331,0.0076071783,0.00089396705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008590142,0.0011566255,0.0007498897,0.008377481,0.00082687906,0.0019753838,0.0019164608,0.001101395,0.0044692964],"category_scores_gemma":[0.030409602,0.00050665304,0.0010088898,0.0052698757,0.00028173253,0.0014482774,0.0028048414,0.0009529495,0.006131353],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013795474,0.00056928804,0.059683498,0.004876087,0.00027109528,0.0002699883,0.0011066047,0.003986473,0.0041863266,0.0020831914,0.726298,0.19529],"study_design_scores_gemma":[0.0014617034,0.00036297133,0.08807499,0.0010689963,0.00027314114,0.0004588818,0.0011310508,0.041202955,0.012359457,0.0069988077,0.84633195,0.0002750538],"about_ca_topic_score_codex":0.010222284,"about_ca_topic_score_gemma":0.034193024,"teacher_disagreement_score":0.010222284,"about_ca_system_score_codex":0.0009965897,"about_ca_system_score_gemma":0.0021582055,"threshold_uncertainty_score":0.045429587},"labels":[],"label_agreement":null},{"id":"W4393548574","doi":"10.5281/zenodo.4010209","title":"Conclusion Stability for Natural Language Based Mining of Design Discussions: Dataset","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Natural (archaeology); Stability (learning theory); Computer science; Natural language processing; Geography; Archaeology; Machine learning","score_opus":0.051186950147725105,"score_gpt":0.2814128574786919,"score_spread":0.23022590733096676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393548574","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025668903,0.00010952341,0.0006461876,0.00009996872,0.000032963606,0.00007715641,0.9935154,0.0015158366,0.0014360552],"genre_scores_gemma":[0.002408253,0.000038401024,0.0015219335,0.000045357676,0.000007211466,0.00023370577,0.99484706,0.000080791906,0.0008172823],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981401,0.0003064038,0.00020827359,0.00054305256,0.00060314545,0.00019901239],"domain_scores_gemma":[0.99582267,0.0012935098,0.00042817724,0.0011965464,0.0009331204,0.00032596503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015562988,0.0024430638,0.0008153648,0.0038139848,0.0009672919,0.0017884022,0.003095758,0.0024622635,0.024724292],"category_scores_gemma":[0.0072985655,0.0005336661,0.0015398463,0.0032296649,0.00051596604,0.0010793319,0.0016981548,0.0017001064,0.03420071],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034284423,0.00021326604,0.004890903,0.0017606887,0.00008276311,0.000091797745,0.00011161492,0.0019022494,0.00089023204,0.0012978079,0.9747677,0.013648031],"study_design_scores_gemma":[0.0006790787,0.0002002036,0.02292413,0.00049243175,0.000119146156,0.00028832504,0.00033278824,0.0066134166,0.004490438,0.0033295269,0.96043366,0.00009684722],"about_ca_topic_score_codex":0.0107736355,"about_ca_topic_score_gemma":0.019977808,"teacher_disagreement_score":0.024724292,"about_ca_system_score_codex":0.0015739466,"about_ca_system_score_gemma":0.0024882047,"threshold_uncertainty_score":0.08271104},"labels":[],"label_agreement":null},{"id":"W4393569314","doi":"10.5281/zenodo.7549217","title":"Bugsplainer: Explaining Software Bugs Leveraging Code Structures in Neural Machine Translation","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Programming language; Software bug; Code (set theory); Translation (biology); Machine translation; Software; Software engineering; Operating system; Artificial intelligence; Biology","score_opus":0.04739010808059809,"score_gpt":0.26718827025099434,"score_spread":0.21979816217039624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393569314","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1693719,0.008616614,0.08887042,0.0039379653,0.0008492367,0.0010515711,0.61871004,0.095821,0.012771297],"genre_scores_gemma":[0.09175816,0.0009233145,0.07584169,0.0005565897,0.00007377857,0.00054669793,0.8254523,0.00085903006,0.0039884658],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99893445,0.000311351,0.00009348122,0.00037152303,0.00020740427,0.00008179379],"domain_scores_gemma":[0.9971776,0.0015738528,0.00022613477,0.0005597626,0.00038002103,0.00008260444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001290453,0.0026408175,0.00058421434,0.0040319264,0.00082811073,0.0009520031,0.0025376568,0.0029114624,0.0057114563],"category_scores_gemma":[0.0078048846,0.00054294,0.0016886869,0.0024377923,0.0005995636,0.0018640662,0.0017665505,0.0020572324,0.005269921],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008883722,0.0007952684,0.026861744,0.003501658,0.00037937044,0.0018917955,0.0006911488,0.035745554,0.0062118084,0.0060257153,0.6947535,0.22225405],"study_design_scores_gemma":[0.0015302217,0.0006115472,0.029330472,0.0006528098,0.00041365225,0.002289394,0.00074480916,0.5456747,0.01912954,0.034479473,0.3648967,0.00024671183],"about_ca_topic_score_codex":0.023586364,"about_ca_topic_score_gemma":0.06813236,"teacher_disagreement_score":0.023586364,"about_ca_system_score_codex":0.0015925461,"about_ca_system_score_gemma":0.0018146063,"threshold_uncertainty_score":0.046898127},"labels":[],"label_agreement":null},{"id":"W4393585945","doi":"10.5281/zenodo.4719161","title":"Automated Support for Searching and Selecting Evidence in Software Engineering: A Cross-domain Systematic Mapping","year":2021,"lang":"en","type":"dataset","venue":"Figshare","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Domain (mathematical analysis); Software engineering; Computer science; Software; Data mining; Data science; Programming language; Mathematics","score_opus":0.06964552652961997,"score_gpt":0.3342136661265946,"score_spread":0.2645681395969746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393585945","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002221936,0.005913953,0.001953642,0.0004912464,0.00004316153,0.002714389,0.9842552,0.00047663937,0.0019299312],"genre_scores_gemma":[0.022596024,0.01044054,0.05001122,0.0014144792,0.00006465174,0.051578574,0.8610576,0.0005898016,0.0022470928],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9740303,0.008046336,0.012122961,0.0024489872,0.0027749932,0.0005763994],"domain_scores_gemma":[0.8191742,0.13246013,0.016827174,0.012224051,0.017585648,0.0017287695],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01824543,0.0016499229,0.004599412,0.02802405,0.0018379653,0.0038842312,0.0029019006,0.0027194808,0.05830337],"category_scores_gemma":[0.17324431,0.0015809257,0.0056848694,0.024324192,0.0006888138,0.002728943,0.0071801413,0.001992889,0.009311719],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002369921,0.0001749013,0.011262203,0.5079411,0.008712809,0.00043162183,0.0018268639,0.0016797932,0.0020295382,0.0039734095,0.3629417,0.09665603],"study_design_scores_gemma":[0.009902869,0.00039286035,0.042542044,0.16591561,0.02137222,0.0005960391,0.0012323569,0.0011720201,0.0022732136,0.009381855,0.74483836,0.0003805357],"about_ca_topic_score_codex":0.011249556,"about_ca_topic_score_gemma":0.04418327,"teacher_disagreement_score":0.98175454,"about_ca_system_score_codex":0.003242867,"about_ca_system_score_gemma":0.013367877,"threshold_uncertainty_score":0.19504422},"labels":[],"label_agreement":null},{"id":"W4393614072","doi":"10.5281/zenodo.7814123","title":"Dataset for \"Deep Dive into the Verifiability of Code Generated by GitHub Copilot\"","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code (set theory); Computer science; Programming language","score_opus":0.03581037201718835,"score_gpt":0.2719287974016264,"score_spread":0.23611842538443806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393614072","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011352075,0.00022894816,0.00041454585,0.0001642958,0.00007286842,0.000040092367,0.99182093,0.0043271463,0.0017960444],"genre_scores_gemma":[0.0010131922,0.000047536323,0.0005606423,0.0000744467,0.0000067076576,0.00007108836,0.99762946,0.00018612163,0.00041085508],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979746,0.00031855426,0.0002292837,0.0004827912,0.0007400211,0.00025488058],"domain_scores_gemma":[0.99628323,0.0012866624,0.00029490024,0.0011118692,0.00076217175,0.0002611859],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0012905255,0.004238026,0.001445764,0.0033783077,0.0012495558,0.0019408885,0.0041390755,0.0039679916,0.028012957],"category_scores_gemma":[0.0076814955,0.00075843406,0.0017876669,0.0033609984,0.0008768523,0.0013777339,0.002682678,0.002652802,0.049461327],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015396408,0.00006246536,0.0007981038,0.0010724752,0.000048961683,0.0001024657,0.00003182611,0.0009607631,0.0004560722,0.0007626752,0.9917018,0.0038485446],"study_design_scores_gemma":[0.0010706183,0.00010762543,0.0074791084,0.00061009236,0.000093481,0.0003898773,0.00012060778,0.0062642586,0.0030623206,0.0050821495,0.97562814,0.00009174235],"about_ca_topic_score_codex":0.022203594,"about_ca_topic_score_gemma":0.046191867,"teacher_disagreement_score":0.9987095,"about_ca_system_score_codex":0.0018999478,"about_ca_system_score_gemma":0.0027607044,"threshold_uncertainty_score":0.09371269},"labels":[],"label_agreement":null},{"id":"W4393616364","doi":"10.5281/zenodo.7226609","title":"Integrating Visual Aids to Enhance the Code Reviewer Selection Process Replication Package","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Replication (statistics); Selection (genetic algorithm); Computer science; Code (set theory); R package; Process (computing); Biology; Programming language; Artificial intelligence; Virology","score_opus":0.02873166098799828,"score_gpt":0.3243248768035528,"score_spread":0.29559321581555453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393616364","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0750743,0.002285375,0.03719515,0.0017951052,0.00067488523,0.0026518703,0.8365082,0.034049146,0.009765935],"genre_scores_gemma":[0.057513013,0.00053192436,0.09328509,0.00040733954,0.00012740826,0.0035719653,0.8397211,0.00080859597,0.004033589],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99411523,0.0023102087,0.0010587957,0.001177652,0.0010382434,0.00029987484],"domain_scores_gemma":[0.9693554,0.012689481,0.0029916088,0.006462331,0.0076071783,0.00089396705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008590142,0.0011566255,0.0007498897,0.008377481,0.00082687906,0.0019753838,0.0019164608,0.001101395,0.0044692964],"category_scores_gemma":[0.030409602,0.00050665304,0.0010088898,0.0052698757,0.00028173253,0.0014482774,0.0028048414,0.0009529495,0.006131353],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013795474,0.00056928804,0.059683498,0.004876087,0.00027109528,0.0002699883,0.0011066047,0.003986473,0.0041863266,0.0020831914,0.726298,0.19529],"study_design_scores_gemma":[0.0014617034,0.00036297133,0.08807499,0.0010689963,0.00027314114,0.0004588818,0.0011310508,0.041202955,0.012359457,0.0069988077,0.84633195,0.0002750538],"about_ca_topic_score_codex":0.010222284,"about_ca_topic_score_gemma":0.034193024,"teacher_disagreement_score":0.010222284,"about_ca_system_score_codex":0.0009965897,"about_ca_system_score_gemma":0.0021582055,"threshold_uncertainty_score":0.045429587},"labels":[],"label_agreement":null},{"id":"W4393623596","doi":"10.5281/zenodo.5173099","title":"Mapping breakpoint types: an exploratory study","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Breakpoint; Computer science; Cartography; Geography; Biology; Genetics; Chromosomal translocation","score_opus":0.04548804132398894,"score_gpt":0.26224738158633415,"score_spread":0.21675934026234522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393623596","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010692281,0.0007787484,0.0011345071,0.00023170409,0.000045280372,0.00014532731,0.9834765,0.00058670883,0.002908831],"genre_scores_gemma":[0.007047801,0.00033947744,0.0026555224,0.00012518214,0.000014543493,0.0005419131,0.9877202,0.0001579149,0.001397423],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99656653,0.00077895034,0.000757929,0.00077930256,0.0007932627,0.0003239967],"domain_scores_gemma":[0.987189,0.0061129457,0.0014985856,0.0018361314,0.0028466117,0.0005167323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027720477,0.0010031649,0.00082368514,0.011538304,0.0011008583,0.002233306,0.0016408009,0.0015714302,0.01271384],"category_scores_gemma":[0.015262733,0.0004199567,0.00081950944,0.013526053,0.0004660959,0.0016045535,0.0025431556,0.0012574417,0.0155073395],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070925517,0.00030143157,0.040027086,0.007996334,0.00015657635,0.0007432704,0.0013820379,0.00085544225,0.0024202329,0.004236168,0.89678925,0.044382896],"study_design_scores_gemma":[0.0002454841,0.00007293747,0.041429453,0.001294257,0.000081775885,0.00055102864,0.0014738826,0.001085002,0.0022281755,0.0017674603,0.9497087,0.00006169183],"about_ca_topic_score_codex":0.0098276185,"about_ca_topic_score_gemma":0.021743547,"teacher_disagreement_score":0.01271384,"about_ca_system_score_codex":0.0016485866,"about_ca_system_score_gemma":0.0020737168,"threshold_uncertainty_score":0.042532086},"labels":[],"label_agreement":null},{"id":"W4393650234","doi":"10.5281/zenodo.4010208","title":"Conclusion Stability for Natural Language Based Mining of Design Discussions: Dataset","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Stability (learning theory); Natural (archaeology); Computer science; Natural language processing; Geography; Machine learning; Archaeology","score_opus":0.051186950147725105,"score_gpt":0.2814128574786919,"score_spread":0.23022590733096676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393650234","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025668903,0.00010952341,0.0006461876,0.00009996872,0.000032963606,0.00007715641,0.9935154,0.0015158366,0.0014360552],"genre_scores_gemma":[0.002408253,0.000038401024,0.0015219335,0.000045357676,0.000007211466,0.00023370577,0.99484706,0.000080791906,0.0008172823],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981401,0.0003064038,0.00020827359,0.00054305256,0.00060314545,0.00019901239],"domain_scores_gemma":[0.99582267,0.0012935098,0.00042817724,0.0011965464,0.0009331204,0.00032596503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015562988,0.0024430638,0.0008153648,0.0038139848,0.0009672919,0.0017884022,0.003095758,0.0024622635,0.024724292],"category_scores_gemma":[0.0072985655,0.0005336661,0.0015398463,0.0032296649,0.00051596604,0.0010793319,0.0016981548,0.0017001064,0.03420071],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034284423,0.00021326604,0.004890903,0.0017606887,0.00008276311,0.000091797745,0.00011161492,0.0019022494,0.00089023204,0.0012978079,0.9747677,0.013648031],"study_design_scores_gemma":[0.0006790787,0.0002002036,0.02292413,0.00049243175,0.000119146156,0.00028832504,0.00033278824,0.0066134166,0.004490438,0.0033295269,0.96043366,0.00009684722],"about_ca_topic_score_codex":0.0107736355,"about_ca_topic_score_gemma":0.019977808,"teacher_disagreement_score":0.024724292,"about_ca_system_score_codex":0.0015739466,"about_ca_system_score_gemma":0.0024882047,"threshold_uncertainty_score":0.08271104},"labels":[],"label_agreement":null},{"id":"W4393653611","doi":"10.5281/zenodo.4019184","title":"On the Untriviality of Trivial Packages: An Empirical Study of npm JavaScript Packages","year":2019,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Concordia University; Queen's University","funders":"","keywords":"JavaScript; Computer science; Programming language; Empirical research; Mathematics; Statistics","score_opus":0.07002625002808652,"score_gpt":0.3073490316570811,"score_spread":0.2373227816289946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393653611","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.996357,0.00017279982,0.0014996708,0.00016493045,0.000006045893,0.000020931571,0.0001792348,0.000028892675,0.0015704116],"genre_scores_gemma":[0.9972953,0.00011952348,0.0015539774,0.00007096155,0.000014065101,0.00003260493,0.00041122324,0.000060916354,0.00044138118],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99255884,0.0034527287,0.0004311219,0.0012667003,0.0018100953,0.00048050564],"domain_scores_gemma":[0.8578404,0.103504926,0.02170451,0.005285066,0.007900701,0.0037643653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009197252,0.00035520268,0.00043744262,0.0037209962,0.0014544812,0.0022193613,0.0012272368,0.0009919315,0.002079457],"category_scores_gemma":[0.07358988,0.0003784176,0.00036919207,0.0037434178,0.0025693316,0.0065027014,0.0020985813,0.0017167601,0.0007787911],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014772265,0.00025864862,0.9412553,0.00028603128,0.00008605195,0.0005170327,0.024798628,0.0004454803,0.001307128,0.0027501918,0.0025343965,0.02561336],"study_design_scores_gemma":[0.000017853532,0.00022229114,0.9339911,0.00020753313,0.00006479376,0.0012834962,0.039466586,0.010461834,0.00092813815,0.0025085614,0.010788181,0.000059612827],"about_ca_topic_score_codex":0.0022392215,"about_ca_topic_score_gemma":0.003329772,"teacher_disagreement_score":0.009197252,"about_ca_system_score_codex":0.0006613132,"about_ca_system_score_gemma":0.00054858014,"threshold_uncertainty_score":0.04864031},"labels":[],"label_agreement":null},{"id":"W4393699540","doi":"10.5281/zenodo.3633081","title":"Empirical Study of the Relationship between Design Patterns and Code Smells","year":2020,"lang":"en","type":"dataset","venue":"Figshare","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code smell; Code (set theory); Computer science; Empirical research; Programming language; Mathematics; Statistics; Software quality; Software","score_opus":0.2551029807026982,"score_gpt":0.368570685920784,"score_spread":0.1134677052180858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393699540","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99756694,0.0000880562,0.0015828192,0.000045574296,0.0000036558163,0.000046561716,0.00023763839,0.000018816223,0.0004099898],"genre_scores_gemma":[0.9973642,0.000036213292,0.001932411,0.000012680267,0.0000053694844,0.0000774254,0.0004113448,0.000010744087,0.0001495534],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98392165,0.006164618,0.0024838801,0.0024466505,0.004505711,0.00047750684],"domain_scores_gemma":[0.5153909,0.3780537,0.06246068,0.019530034,0.021944126,0.0026205776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013941171,0.000426702,0.00038211377,0.0032272006,0.0005287691,0.0011646765,0.0007461306,0.0007803991,0.0011030644],"category_scores_gemma":[0.15973903,0.00039194073,0.00060123595,0.0033123174,0.0012867956,0.0020474282,0.0014741396,0.0011697196,0.0002469968],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016617906,0.00024037855,0.9833182,0.00016108052,0.00015299307,0.000121125064,0.0021035175,0.0006285213,0.0008531946,0.00017639415,0.00018739027,0.011890986],"study_design_scores_gemma":[0.000024977699,0.00050851854,0.98745406,0.00005193799,0.00007573288,0.00036098552,0.0024572418,0.005980544,0.0015672835,0.0004165444,0.0010737358,0.000028329201],"about_ca_topic_score_codex":0.0011436979,"about_ca_topic_score_gemma":0.0015897062,"teacher_disagreement_score":0.013941171,"about_ca_system_score_codex":0.00064455485,"about_ca_system_score_gemma":0.00061497494,"threshold_uncertainty_score":0.0737288},"labels":[],"label_agreement":null},{"id":"W4393720210","doi":"10.5281/zenodo.5066797","title":"Automated Support for Searching and Selecting Evidence in Software Engineering: A Cross-domain Systematic Mapping","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Domain (mathematical analysis); Computer science; Software engineering; Software; Data mining; Programming language; Mathematics","score_opus":0.043030932149270076,"score_gpt":0.28234802919294694,"score_spread":0.23931709704367687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393720210","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000825194,0.0006756112,0.0007688434,0.00011559981,0.000012978106,0.00026998867,0.99636567,0.00025948565,0.0007066209],"genre_scores_gemma":[0.002866432,0.00068296643,0.0071046543,0.00013435805,0.000007761225,0.002778175,0.985786,0.00009868092,0.00054095866],"study_design_codex":"not_applicable","study_design_gemma":"systematic_review","domain_scores_codex":[0.9895392,0.003036206,0.0038516286,0.0016168947,0.001575437,0.00038075598],"domain_scores_gemma":[0.9428221,0.038210917,0.0053106467,0.00602909,0.006621509,0.0010056336],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009594801,0.0015797489,0.0023590357,0.017956024,0.0013391785,0.0030455652,0.0024646688,0.0023773443,0.03598304],"category_scores_gemma":[0.06850581,0.0009766918,0.0023440744,0.018012518,0.00065424095,0.0015197651,0.0049121208,0.0015329928,0.013917737],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013240663,0.00019417175,0.0135334525,0.085146494,0.0018267967,0.00045289818,0.0010589621,0.002605608,0.0034871644,0.0056260284,0.81244767,0.07229682],"study_design_scores_gemma":[0.001738665,0.00012860779,0.02600781,0.012028879,0.0015363579,0.0003827039,0.0005525008,0.001062432,0.0023572342,0.0053956625,0.9486727,0.00013652939],"about_ca_topic_score_codex":0.014939926,"about_ca_topic_score_gemma":0.040987976,"teacher_disagreement_score":0.9904052,"about_ca_system_score_codex":0.002695211,"about_ca_system_score_gemma":0.0085622,"threshold_uncertainty_score":0.120375335},"labels":[],"label_agreement":null},{"id":"W4393743078","doi":"10.5281/zenodo.7479227","title":"Dataset for a systematic literature review on source code similarity measurement and clone detection: techniques, applications, and challenges","year":2022,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Similarity (geometry); Computer science; Code (set theory); Source code; clone (Java method); Information retrieval; Data mining; Programming language; Artificial intelligence; Biology; Genetics; DNA","score_opus":0.05538046540992728,"score_gpt":0.26851299974022763,"score_spread":0.21313253433030035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393743078","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00030182014,0.001187802,0.00019312675,0.00019441254,0.000035372046,0.00037930068,0.9970065,0.0000879975,0.0006137292],"genre_scores_gemma":[0.003433634,0.0033714678,0.0035231386,0.00074325065,0.000050132803,0.012295426,0.97459453,0.00013380854,0.001854663],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9915598,0.0017913358,0.0037598177,0.0009570392,0.001462725,0.0004692568],"domain_scores_gemma":[0.955836,0.027848534,0.005981456,0.0023645256,0.007061802,0.0009076877],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0056060613,0.0018753535,0.0036653958,0.01996318,0.0011857033,0.0025365248,0.0020834366,0.0022663893,0.120073766],"category_scores_gemma":[0.05338473,0.0009242213,0.0031032737,0.025464509,0.00059161475,0.0018377021,0.0033812053,0.002139964,0.020907894],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000914685,0.00009136157,0.0021822327,0.23254873,0.0011933314,0.00027684917,0.00037713422,0.0006549433,0.00076482544,0.0021220872,0.7343456,0.02452825],"study_design_scores_gemma":[0.0024571996,0.00013620238,0.013151269,0.057914585,0.0019246074,0.0003689,0.00058387243,0.0002228545,0.00061747665,0.0029495945,0.9195285,0.00014502539],"about_ca_topic_score_codex":0.012055243,"about_ca_topic_score_gemma":0.032085873,"teacher_disagreement_score":0.99439394,"about_ca_system_score_codex":0.0030476262,"about_ca_system_score_gemma":0.011184162,"threshold_uncertainty_score":0.4016868},"labels":[],"label_agreement":null},{"id":"W4393776068","doi":"10.5281/zenodo.4732665","title":"Dataset of the Paper \"System and Software Processes in Practice: Insights from Chinese Industry\"","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Software; Computer science; Data science; Software engineering; Business; Operating system","score_opus":0.018096455100356325,"score_gpt":0.25348106174621415,"score_spread":0.23538460664585784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393776068","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012746669,0.00010984198,0.00010476868,0.00014061618,0.000042374682,0.00003371732,0.9968671,0.00017096862,0.0012558837],"genre_scores_gemma":[0.00160596,0.00005844392,0.00016381861,0.000039810955,0.0000080733835,0.0001321048,0.997288,0.000015274894,0.00068857236],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984275,0.00029663614,0.00019638434,0.00028586018,0.00046116117,0.000332366],"domain_scores_gemma":[0.99648094,0.0009074121,0.00029799851,0.00066586654,0.0012097702,0.00043796675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013093669,0.0015537152,0.0012160839,0.004692449,0.0010339366,0.0014676355,0.0028276502,0.002081132,0.029218992],"category_scores_gemma":[0.0062577124,0.00039995988,0.0010486266,0.008588038,0.0005095114,0.0010532095,0.0024396472,0.0013308186,0.028539952],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015992373,0.00009604575,0.003971437,0.0011545329,0.000048031914,0.000097229626,0.00011147687,0.0006202084,0.00019364324,0.00094172155,0.98509824,0.0075075505],"study_design_scores_gemma":[0.00030418613,0.00006678649,0.051059507,0.0006001582,0.00009539391,0.00014926751,0.00068154675,0.0014524579,0.0007352346,0.0015733086,0.9432141,0.00006795633],"about_ca_topic_score_codex":0.049589034,"about_ca_topic_score_gemma":0.07377536,"teacher_disagreement_score":0.049589034,"about_ca_system_score_codex":0.0024438435,"about_ca_system_score_gemma":0.0039641922,"threshold_uncertainty_score":0.098600805},"labels":[],"label_agreement":null},{"id":"W4393846738","doi":"10.5281/zenodo.7812451","title":"Dataset for \"Deep Dive into the Verifiability of Code Generated by GitHub Copilot\"","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code (set theory); Computer science; Programming language","score_opus":0.03581037201718835,"score_gpt":0.2719287974016264,"score_spread":0.23611842538443806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393846738","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011352075,0.00022894816,0.00041454585,0.0001642958,0.00007286842,0.000040092367,0.99182093,0.0043271463,0.0017960444],"genre_scores_gemma":[0.0010131922,0.000047536323,0.0005606423,0.0000744467,0.0000067076576,0.00007108836,0.99762946,0.00018612163,0.00041085508],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979746,0.00031855426,0.0002292837,0.0004827912,0.0007400211,0.00025488058],"domain_scores_gemma":[0.99628323,0.0012866624,0.00029490024,0.0011118692,0.00076217175,0.0002611859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012905255,0.004238026,0.001445764,0.0033783077,0.0012495558,0.0019408885,0.0041390755,0.0039679916,0.028012957],"category_scores_gemma":[0.0076814955,0.00075843406,0.0017876669,0.0033609984,0.0008768523,0.0013777339,0.002682678,0.002652802,0.049461327],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015396408,0.00006246536,0.0007981038,0.0010724752,0.000048961683,0.0001024657,0.00003182611,0.0009607631,0.0004560722,0.0007626752,0.9917018,0.0038485446],"study_design_scores_gemma":[0.0010706183,0.00010762543,0.0074791084,0.00061009236,0.000093481,0.0003898773,0.00012060778,0.0062642586,0.0030623206,0.0050821495,0.97562814,0.00009174235],"about_ca_topic_score_codex":0.022203594,"about_ca_topic_score_gemma":0.046191867,"teacher_disagreement_score":0.028012957,"about_ca_system_score_codex":0.0018999478,"about_ca_system_score_gemma":0.0027607044,"threshold_uncertainty_score":0.09371269},"labels":[],"label_agreement":null},{"id":"W4393851745","doi":"10.1145/3650142.3650148","title":"Column: Is Theory (Still) Welcome in Software Engineering Research?","year":2024,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Column (typography); Computer science; Software; Software engineering; Programming language; Telecommunications","score_opus":0.03368683316757169,"score_gpt":0.3006107158940571,"score_spread":0.2669238827264854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393851745","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005894958,0.0027455206,0.00046075997,0.58844024,0.35203674,0.00011606942,0.001732437,0.0007515641,0.05312718],"genre_scores_gemma":[0.007266538,0.003483603,0.00085907127,0.44363514,0.22521786,0.000322996,0.002380103,0.0013275105,0.3155073],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.996436,0.0006630332,0.00029550938,0.00058268855,0.0014964087,0.00052627514],"domain_scores_gemma":[0.9656992,0.008598891,0.0018234215,0.0013407768,0.012785774,0.009752033],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0036048822,0.0016107952,0.001938908,0.0015588588,0.0047009927,0.0139641445,0.0016993452,0.008539514,0.29995587],"category_scores_gemma":[0.023967499,0.000804402,0.0012105538,0.0019045414,0.0026364883,0.006555536,0.002505611,0.010152182,0.16325505],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034870012,0.0000102820395,0.000056798715,0.00003422552,0.0000029943649,0.000012586158,0.000010816864,0.0000076083356,0.00004362657,0.00025207436,0.9970483,0.0024858145],"study_design_scores_gemma":[0.00008819236,0.00006366632,0.0010678164,0.00028197208,0.000014886805,0.00008499433,0.00026802652,0.00020854844,0.00029176052,0.0022202143,0.99538237,0.00002749632],"about_ca_topic_score_codex":0.002436874,"about_ca_topic_score_gemma":0.0044046943,"teacher_disagreement_score":0.9963951,"about_ca_system_score_codex":0.0028683094,"about_ca_system_score_gemma":0.0056099594,"threshold_uncertainty_score":0.998528},"labels":[],"label_agreement":null},{"id":"W4393886095","doi":"10.5281/zenodo.4726379","title":"Dataset of the Paper \"System and Software Processes in Practice: Insights from Chinese Industry\"","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Software; Computer science; Software engineering; Data science; Programming language","score_opus":0.018096455100356325,"score_gpt":0.25348106174621415,"score_spread":0.23538460664585784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393886095","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012746669,0.00010984198,0.00010476868,0.00014061618,0.000042374682,0.00003371732,0.9968671,0.00017096862,0.0012558837],"genre_scores_gemma":[0.00160596,0.00005844392,0.00016381861,0.000039810955,0.0000080733835,0.0001321048,0.997288,0.000015274894,0.00068857236],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984275,0.00029663614,0.00019638434,0.00028586018,0.00046116117,0.000332366],"domain_scores_gemma":[0.99648094,0.0009074121,0.00029799851,0.00066586654,0.0012097702,0.00043796675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013093669,0.0015537152,0.0012160839,0.004692449,0.0010339366,0.0014676355,0.0028276502,0.002081132,0.029218992],"category_scores_gemma":[0.0062577124,0.00039995988,0.0010486266,0.008588038,0.0005095114,0.0010532095,0.0024396472,0.0013308186,0.028539952],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015992373,0.00009604575,0.003971437,0.0011545329,0.000048031914,0.000097229626,0.00011147687,0.0006202084,0.00019364324,0.00094172155,0.98509824,0.0075075505],"study_design_scores_gemma":[0.00030418613,0.00006678649,0.051059507,0.0006001582,0.00009539391,0.00014926751,0.00068154675,0.0014524579,0.0007352346,0.0015733086,0.9432141,0.00006795633],"about_ca_topic_score_codex":0.049589034,"about_ca_topic_score_gemma":0.07377536,"teacher_disagreement_score":0.049589034,"about_ca_system_score_codex":0.0024438435,"about_ca_system_score_gemma":0.0039641922,"threshold_uncertainty_score":0.098600805},"labels":[],"label_agreement":null},{"id":"W4393924710","doi":"10.48550/arxiv.2404.01470","title":"Measuring the Redundancy of Information from a Source Failure Perspective","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Redundancy (engineering); Perspective (graphical); Computer science; Reliability engineering; Information retrieval; Engineering; Artificial intelligence","score_opus":0.04600702507292608,"score_gpt":0.18307209046752596,"score_spread":0.13706506539459987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393924710","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.083604135,0.0013423106,0.9052173,0.00072202453,0.000101805024,0.00008754157,0.00048452755,0.0004018329,0.008038548],"genre_scores_gemma":[0.86741775,0.000829741,0.12864831,0.0002520681,0.00026728614,0.0002074771,0.000598287,0.00014618985,0.001632871],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99004525,0.0023710264,0.0007717361,0.0015378697,0.004774546,0.0004994949],"domain_scores_gemma":[0.9684221,0.017596373,0.004184674,0.005591183,0.0031985838,0.0010070893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00662512,0.0012975754,0.0018879072,0.006695378,0.0010795081,0.003148738,0.0022353746,0.0018725903,0.0022222828],"category_scores_gemma":[0.029800588,0.00072405214,0.0013091872,0.0033254973,0.0047613443,0.011085242,0.0046680947,0.0022126934,0.00049797975],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007163875,0.00020222965,0.013564792,0.00093971693,0.0005067339,0.00083534577,0.0014989792,0.17158405,0.024097903,0.68056464,0.0025247233,0.10296443],"study_design_scores_gemma":[0.000062082836,0.00048468096,0.006255458,0.0002023183,0.00025942337,0.0013022968,0.00046903206,0.2719392,0.017805448,0.69323385,0.007793913,0.0001922869],"about_ca_topic_score_codex":0.00051755115,"about_ca_topic_score_gemma":0.00022937436,"teacher_disagreement_score":0.006695378,"about_ca_system_score_codex":0.0016380909,"about_ca_system_score_gemma":0.0009454979,"threshold_uncertainty_score":0.0350374},"labels":[],"label_agreement":null},{"id":"W4394046374","doi":"10.5281/zenodo.5066796","title":"Automated Support for Searching and Selecting Evidence in Software Engineering: A Cross-domain Systematic Mapping","year":2021,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Software engineering; Computer science; Domain (mathematical analysis); Software; Data mining; Data science; Information retrieval; Programming language; Mathematics","score_opus":0.043030932149270076,"score_gpt":0.28234802919294694,"score_spread":0.23931709704367687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394046374","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000825194,0.0006756112,0.0007688434,0.00011559981,0.000012978106,0.00026998867,0.99636567,0.00025948565,0.0007066209],"genre_scores_gemma":[0.002866432,0.00068296643,0.0071046543,0.00013435805,0.000007761225,0.002778175,0.985786,0.00009868092,0.00054095866],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9895392,0.003036206,0.0038516286,0.0016168947,0.001575437,0.00038075598],"domain_scores_gemma":[0.9428221,0.038210917,0.0053106467,0.00602909,0.006621509,0.0010056336],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009594801,0.0015797489,0.0023590357,0.017956024,0.0013391785,0.0030455652,0.0024646688,0.0023773443,0.03598304],"category_scores_gemma":[0.06850581,0.0009766918,0.0023440744,0.018012518,0.00065424095,0.0015197651,0.0049121208,0.0015329928,0.013917737],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013240663,0.00019417175,0.0135334525,0.085146494,0.0018267967,0.00045289818,0.0010589621,0.002605608,0.0034871644,0.0056260284,0.81244767,0.07229682],"study_design_scores_gemma":[0.001738665,0.00012860779,0.02600781,0.012028879,0.0015363579,0.0003827039,0.0005525008,0.001062432,0.0023572342,0.0053956625,0.9486727,0.00013652939],"about_ca_topic_score_codex":0.014939926,"about_ca_topic_score_gemma":0.040987976,"teacher_disagreement_score":0.03598304,"about_ca_system_score_codex":0.002695211,"about_ca_system_score_gemma":0.0085622,"threshold_uncertainty_score":0.120375335},"labels":[],"label_agreement":null},{"id":"W4394051971","doi":"10.5281/zenodo.7636742","title":"Bugsplainer: Explaining Software Bugs Leveraging Code Structures in Neural Machine Translation","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Machine translation; Code (set theory); Programming language; Software; Translation (biology); Software bug; Artificial intelligence; Software engineering; Biology","score_opus":0.04739010808059809,"score_gpt":0.26718827025099434,"score_spread":0.21979816217039624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394051971","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1693719,0.008616614,0.08887042,0.0039379653,0.0008492367,0.0010515711,0.61871004,0.095821,0.012771297],"genre_scores_gemma":[0.09175816,0.0009233145,0.07584169,0.0005565897,0.00007377857,0.00054669793,0.8254523,0.00085903006,0.0039884658],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99893445,0.000311351,0.00009348122,0.00037152303,0.00020740427,0.00008179379],"domain_scores_gemma":[0.9971776,0.0015738528,0.00022613477,0.0005597626,0.00038002103,0.00008260444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001290453,0.0026408175,0.00058421434,0.0040319264,0.00082811073,0.0009520031,0.0025376568,0.0029114624,0.0057114563],"category_scores_gemma":[0.0078048846,0.00054294,0.0016886869,0.0024377923,0.0005995636,0.0018640662,0.0017665505,0.0020572324,0.005269921],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008883722,0.0007952684,0.026861744,0.003501658,0.00037937044,0.0018917955,0.0006911488,0.035745554,0.0062118084,0.0060257153,0.6947535,0.22225405],"study_design_scores_gemma":[0.0015302217,0.0006115472,0.029330472,0.0006528098,0.00041365225,0.002289394,0.00074480916,0.5456747,0.01912954,0.034479473,0.3648967,0.00024671183],"about_ca_topic_score_codex":0.023586364,"about_ca_topic_score_gemma":0.06813236,"teacher_disagreement_score":0.023586364,"about_ca_system_score_codex":0.0015925461,"about_ca_system_score_gemma":0.0018146063,"threshold_uncertainty_score":0.046898127},"labels":[],"label_agreement":null},{"id":"W4394745299","doi":"10.1145/3597503.3639177","title":"Demystifying and Detecting Misuses of Deep Learning APIs","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; York University","funders":"","keywords":"Computer science; Software; Artificial intelligence; Focus (optics); World Wide Web; Computer security; Programming language","score_opus":0.020221029892897965,"score_gpt":0.27987897484834745,"score_spread":0.25965794495544947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394745299","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7407388,0.0020833446,0.18836723,0.0013861096,0.0002561264,0.00023069727,0.003645263,0.06029737,0.0029949676],"genre_scores_gemma":[0.8965769,0.0005135702,0.0925381,0.00054656336,0.000054822056,0.00020502109,0.006288225,0.0012048787,0.0020720172],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9957444,0.0008278965,0.0004964097,0.0008403063,0.0015888604,0.00050210295],"domain_scores_gemma":[0.98415977,0.005945748,0.002430077,0.0038977258,0.0030078902,0.0005588316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030735706,0.0020206394,0.0010053641,0.004873835,0.000654346,0.0012713058,0.0019119259,0.0010010921,0.0006138146],"category_scores_gemma":[0.022479173,0.00075298356,0.0012332827,0.002136384,0.0010959259,0.0042077526,0.0031652898,0.0024247572,0.0005276647],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009790959,0.00078591396,0.28710675,0.00101434,0.00043508425,0.0027381184,0.0024180857,0.049257774,0.040355682,0.0035428868,0.026349224,0.585017],"study_design_scores_gemma":[0.000037787337,0.0002552905,0.020719785,0.00012387699,0.00012955582,0.00086238154,0.0003913799,0.92897,0.033783264,0.006119401,0.008529237,0.00007813025],"about_ca_topic_score_codex":0.009096368,"about_ca_topic_score_gemma":0.013899922,"teacher_disagreement_score":0.009096368,"about_ca_system_score_codex":0.0010410284,"about_ca_system_score_gemma":0.002214403,"threshold_uncertainty_score":0.018086791},"labels":[],"label_agreement":null},{"id":"W4394746930","doi":"10.1145/3597503.3639110","title":"Towards More Practical Automation of Vulnerability Assessment","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"Fundamental Research Funds for the Central Universities; National Key Research and Development Program of China","keywords":"Computer science; Exploit; Vulnerability (computing); Schedule; Machine learning; Vulnerability assessment; Software; Prioritization; Data science; Artificial intelligence; Risk analysis (engineering); Computer security; Process management; Engineering","score_opus":0.036583347843623895,"score_gpt":0.40381664148310714,"score_spread":0.36723329363948326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394746930","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016105145,0.0005407134,0.95909035,0.0010974755,0.000092414055,0.0002466799,0.0004688736,0.019850932,0.0025075176],"genre_scores_gemma":[0.27951062,0.0006272838,0.7138536,0.00051963597,0.00008172094,0.00031485097,0.0019775904,0.0011576145,0.0019570724],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9882181,0.0056659384,0.0007080235,0.0021837817,0.0027951377,0.00042903514],"domain_scores_gemma":[0.9648729,0.019170005,0.0020482643,0.008631084,0.004370508,0.00090723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009244462,0.0024051974,0.0016059238,0.0029563198,0.0005288205,0.0046304697,0.0032375061,0.0021851012,0.004932845],"category_scores_gemma":[0.045838296,0.0012506567,0.0018057646,0.0013748693,0.0011561718,0.0078039765,0.0065051843,0.005766159,0.0041210963],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003604404,0.0010336974,0.020939833,0.001244659,0.00023035868,0.0005280491,0.0015216062,0.13562962,0.024874007,0.026643395,0.01790623,0.76908815],"study_design_scores_gemma":[0.00004146884,0.00017365692,0.0026588824,0.00025899938,0.000045197896,0.00023177444,0.0002678986,0.9170701,0.008789099,0.052894093,0.017500354,0.00006838027],"about_ca_topic_score_codex":0.0031505227,"about_ca_topic_score_gemma":0.00388434,"teacher_disagreement_score":0.009244462,"about_ca_system_score_codex":0.0011535315,"about_ca_system_score_gemma":0.0028722957,"threshold_uncertainty_score":0.048889935},"labels":[],"label_agreement":null},{"id":"W4394768860","doi":"10.1145/3597503.3639089","title":"Supporting Web-Based API Searches in the IDE Using Signatures","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Task (project management); World Wide Web; Context (archaeology); Application programming interface; Rank (graph theory); Information retrieval; Representation (politics); Programming language","score_opus":0.04268976473481468,"score_gpt":0.3458410115771883,"score_spread":0.3031512468423736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394768860","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21770485,0.00042634385,0.6761637,0.00033125762,0.000059910482,0.0006828199,0.0018499622,0.091948666,0.010832492],"genre_scores_gemma":[0.4359835,0.00030066518,0.54931676,0.0001474074,0.000040521765,0.00021067755,0.0035355596,0.00494419,0.0055207545],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965773,0.0010862642,0.0003492516,0.00053590455,0.0012118978,0.00023937646],"domain_scores_gemma":[0.97869086,0.013559362,0.0016533608,0.0030727456,0.0024169674,0.0006068266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040679984,0.001432963,0.0013362621,0.0032435826,0.00075743004,0.0034610445,0.0011525358,0.0010147628,0.003015937],"category_scores_gemma":[0.021323076,0.0009374088,0.0007035519,0.0016333176,0.00068224734,0.005824329,0.0022516726,0.00092413946,0.003107686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030683996,0.0013478593,0.034109972,0.0011639002,0.00012142029,0.0011636906,0.0043175733,0.007821524,0.12183679,0.0067624697,0.021210607,0.7970757],"study_design_scores_gemma":[0.0005770668,0.0017245982,0.022340937,0.0004135804,0.00021982293,0.0020775127,0.0034161608,0.5218491,0.32728547,0.022176594,0.09744302,0.00047625837],"about_ca_topic_score_codex":0.0021315718,"about_ca_topic_score_gemma":0.00445262,"teacher_disagreement_score":0.0040679984,"about_ca_system_score_codex":0.0005268543,"about_ca_system_score_gemma":0.0015877547,"threshold_uncertainty_score":0.02151388},"labels":[],"label_agreement":null},{"id":"W4394769089","doi":"10.1145/3597503.3639169","title":"The Classics Never Go Out of Style: An Empirical Study of Downgrades from the Bazel Build Technology","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Deliverable; Abandonment (legal); Computer science; sync; Software; Investment (military); Software development; Emerging technologies; Work (physics); Software engineering; Risk analysis (engineering); Systems engineering; Telecommunications; Engineering; Business; Operating system","score_opus":0.03354210998272119,"score_gpt":0.3363972040823928,"score_spread":0.30285509409967165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394769089","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99215746,0.0004411102,0.00056429446,0.00047017616,0.00002359463,0.000030762458,0.00004052961,0.00001570767,0.0062563242],"genre_scores_gemma":[0.99713695,0.00027231633,0.00042464098,0.00027914366,0.000014972414,0.000022045142,0.00008178603,0.00004592359,0.0017222768],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99098325,0.0040389467,0.00076337956,0.0008743573,0.0027783161,0.0005617871],"domain_scores_gemma":[0.8904549,0.052430186,0.030148974,0.009882917,0.012059664,0.0050233984],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013372784,0.00030503992,0.0005319149,0.0019959945,0.0028350158,0.0054198187,0.0017968041,0.0012286233,0.004031921],"category_scores_gemma":[0.09469876,0.0005683968,0.0002759913,0.0019588666,0.005838496,0.0067374287,0.003124535,0.0034583006,0.00093996845],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006347302,0.0009892298,0.38871205,0.00048250007,0.000094488794,0.000718618,0.51978093,0.00024282183,0.0020831514,0.013364228,0.0053724796,0.06752481],"study_design_scores_gemma":[0.00005311084,0.000823567,0.5761111,0.00068846624,0.000085877626,0.001148372,0.36796275,0.001392025,0.0011387791,0.0058760066,0.044565186,0.00015493244],"about_ca_topic_score_codex":0.005149207,"about_ca_topic_score_gemma":0.0051527624,"teacher_disagreement_score":0.013372784,"about_ca_system_score_codex":0.0021576982,"about_ca_system_score_gemma":0.0012874262,"threshold_uncertainty_score":0.07072288},"labels":[],"label_agreement":null},{"id":"W4394769335","doi":"10.1145/3597503.3639166","title":"A First Look at the Inheritance-Induced Redundant Test Execution","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Inheritance (genetic algorithm); Computer science; Code refactoring; Code reuse; Code coverage; Test case; Programming language; Test (biology); Code (set theory); Software quality; Software engineering; Software; Software development; Machine learning","score_opus":0.023573761653133494,"score_gpt":0.2698936586581968,"score_spread":0.24631989700506332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394769335","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8549726,0.00397984,0.12489132,0.0023754567,0.0001421064,0.00018893594,0.0007895368,0.0010379469,0.011622342],"genre_scores_gemma":[0.9616829,0.00090333034,0.034820504,0.00032811554,0.00007423473,0.000050316372,0.0007754198,0.00015620836,0.0012088708],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9938543,0.0013145608,0.0005764152,0.00081748504,0.002909284,0.00052802195],"domain_scores_gemma":[0.93067175,0.042949975,0.011431545,0.005834501,0.008381602,0.0007305482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029594472,0.0004686983,0.0004237403,0.0033766879,0.0006564542,0.0010600691,0.0010172388,0.0006923843,0.0019743594],"category_scores_gemma":[0.031990353,0.00033393523,0.000642054,0.002141663,0.00090780394,0.0017812152,0.0007889648,0.001198634,0.00027271567],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004828025,0.00048450357,0.48603153,0.0009420469,0.00028606682,0.0067892354,0.0029347094,0.012843926,0.043768153,0.020905308,0.0046367804,0.41989496],"study_design_scores_gemma":[0.00009558857,0.0019435053,0.6191229,0.0014307764,0.0008595608,0.037407797,0.0048575765,0.18444753,0.06790053,0.034101732,0.047616314,0.00021620929],"about_ca_topic_score_codex":0.0035433453,"about_ca_topic_score_gemma":0.006795832,"teacher_disagreement_score":0.0035433453,"about_ca_system_score_codex":0.0008379419,"about_ca_system_score_gemma":0.0014675055,"threshold_uncertainty_score":0.015651226},"labels":[],"label_agreement":null},{"id":"W4394769336","doi":"10.1145/3597503.3639212","title":"Combining Structured Static Code Information and Dynamic Symbolic Traces for Software Vulnerability Prediction","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Engineering and Physical Sciences Research Council; Mitacs; National Natural Science Foundation of China","keywords":"Computer science; Symbolic execution; Scalability; Fuzz testing; Overhead (engineering); Code (set theory); Source code; Programming language; Vulnerability (computing); Static analysis; Semantics (computer science); Software; Static program analysis; Benchmark (surveying); Artificial intelligence; Machine learning; Database; Software development; Computer security; Set (abstract data type)","score_opus":0.011120503746424083,"score_gpt":0.27725547842754644,"score_spread":0.26613497468112235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394769336","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28128746,0.0009883747,0.6885914,0.0006510495,0.000044141976,0.00013892267,0.0019810863,0.023261279,0.003056203],"genre_scores_gemma":[0.80250615,0.00036283166,0.1907346,0.00020194864,0.000025071093,0.00009295854,0.0039536497,0.0004177762,0.0017050218],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99936754,0.00012312597,0.00004595185,0.00019994982,0.00018554888,0.000077974495],"domain_scores_gemma":[0.9966793,0.0015251713,0.0005025547,0.0006597332,0.0004821908,0.0001510541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008842477,0.001112884,0.0005060553,0.0040917387,0.00031724284,0.0007492429,0.0011443583,0.00078323385,0.0011615172],"category_scores_gemma":[0.0056534335,0.00047744135,0.0005874651,0.001895301,0.0007329091,0.0027831977,0.0018921872,0.0012636562,0.00056005566],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003146841,0.0005242037,0.06029895,0.0003796939,0.00014983138,0.00027927046,0.0003437312,0.22420402,0.018203251,0.004553714,0.005579132,0.6851694],"study_design_scores_gemma":[0.000009427948,0.00007767458,0.0025273014,0.00002731282,0.000022851009,0.00006526636,0.00005420176,0.98055226,0.0074086436,0.007938198,0.001300369,0.000016462856],"about_ca_topic_score_codex":0.0065873815,"about_ca_topic_score_gemma":0.013045977,"teacher_disagreement_score":0.0065873815,"about_ca_system_score_codex":0.00095144263,"about_ca_system_score_gemma":0.0014094481,"threshold_uncertainty_score":0.013098061},"labels":[],"label_agreement":null},{"id":"W4394866494","doi":"10.48550/arxiv.2404.09384","title":"Generative transformations and patterns in LLM-native approaches for software verification and falsification","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fundação de Amparo à Pesquisa do Estado do Amazonas; Agencia Nacional de Promoción Científica y Tecnológica; International Development Research Centre; Consejo Nacional de Investigaciones Científicas y Técnicas; Coordenação de Aperfeiçoamento de Pessoal de Nível Superior; Universidade Federal do Amazonas; Secretaría de Ciencia y Técnica, Universidad de Buenos Aires; Agencia Nacional de Investigación e Innovación","keywords":"Downstream (manufacturing); Taxonomy (biology); Computer science; Software; Software engineering; Engineering; Programming language; Operations management; Biology; Ecology","score_opus":0.12950047823845698,"score_gpt":0.21804382603729106,"score_spread":0.08854334779883408,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394866494","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013793773,0.00030768942,0.9777945,0.0010090056,0.000031628628,0.00011072319,0.000051080224,0.0011642049,0.005737302],"genre_scores_gemma":[0.35925257,0.00044347983,0.6354099,0.00048499543,0.000041488587,0.00039681746,0.00022986818,0.0006935747,0.0030473387],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9786089,0.01157347,0.0015722058,0.002579027,0.004743744,0.0009227351],"domain_scores_gemma":[0.9637926,0.018342227,0.0023339863,0.012645474,0.0023854696,0.00050018966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015375653,0.000972549,0.0006151124,0.0029295455,0.0020757932,0.0072268466,0.0025876274,0.002460191,0.003434696],"category_scores_gemma":[0.047314156,0.0009998007,0.0014990952,0.0021635806,0.01254221,0.010619504,0.0068029873,0.004102124,0.0010486776],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006139763,0.000067065965,0.0026284847,0.00023499559,0.000023759572,0.00023758634,0.006701265,0.008316627,0.0026454686,0.8939304,0.0009784189,0.08417444],"study_design_scores_gemma":[0.000021911312,0.00007586041,0.0005908975,0.00021067113,0.00003408874,0.00047652415,0.0016218363,0.043617677,0.0073202774,0.9152879,0.0306944,0.000048043912],"about_ca_topic_score_codex":0.0016330205,"about_ca_topic_score_gemma":0.0020570946,"teacher_disagreement_score":0.015375653,"about_ca_system_score_codex":0.0033096797,"about_ca_system_score_gemma":0.004314448,"threshold_uncertainty_score":0.08131522},"labels":[],"label_agreement":null},{"id":"W4395075491","doi":"10.2139/ssrn.4805882","title":"Impact of Methodological Choices on the Analysis of Code Metrics And Maintenance","year":2024,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Code (set theory); Computer science; Risk analysis (engineering); Business; Programming language","score_opus":0.07025245933187206,"score_gpt":0.38033847611651556,"score_spread":0.3100860167846435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395075491","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61100143,0.0079391105,0.34964103,0.013266332,0.0009368695,0.0008430495,0.0019649796,0.00094329944,0.013463841],"genre_scores_gemma":[0.78447765,0.0010225205,0.21024027,0.0010877509,0.00023938438,0.00054006535,0.00065700064,0.0007777082,0.00095762516],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.66933376,0.2435093,0.032810766,0.014322753,0.03730327,0.0027202228],"domain_scores_gemma":[0.103532046,0.7919842,0.033470895,0.04396573,0.02530311,0.001744089],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.23604928,0.00126532,0.0012198577,0.0071119335,0.0019813671,0.0075144423,0.003692697,0.0034351677,0.0022590884],"category_scores_gemma":[0.7000432,0.0009634532,0.0025708107,0.009416836,0.003318019,0.006904257,0.0043012192,0.003762167,0.00046073852],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011397315,0.0021665245,0.2838901,0.0063100983,0.0071877097,0.0010357002,0.009068378,0.0471592,0.02945777,0.1450517,0.006718487,0.45055708],"study_design_scores_gemma":[0.003371872,0.0077886996,0.22203985,0.005555734,0.00906281,0.0019525359,0.009428126,0.2323821,0.07207045,0.40237224,0.03295813,0.0010175683],"about_ca_topic_score_codex":0.0036897997,"about_ca_topic_score_gemma":0.005686519,"teacher_disagreement_score":0.7639507,"about_ca_system_score_codex":0.0046985126,"about_ca_system_score_gemma":0.00663606,"threshold_uncertainty_score":0.94208723},"labels":[],"label_agreement":null},{"id":"W4395097921","doi":"10.1007/978-3-031-57537-2_15","title":"Enhancing Code Security Through Open-Source Large Language Models: A Comparative Study","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Strengths and weaknesses; Code (set theory); Vulnerability (computing); Code review; Source code; Static program analysis; Computer security; Data science; Software engineering; Programming language; Software; Software development","score_opus":0.03903566815580141,"score_gpt":0.3290575166046351,"score_spread":0.2900218484488337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395097921","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9520855,0.0018976543,0.020128535,0.0009789994,0.000033755045,0.000107432614,0.00011243454,0.00039752477,0.024258116],"genre_scores_gemma":[0.98842454,0.00080015016,0.007948113,0.00007997072,0.000009683931,0.000042686454,0.00013526248,0.0001629641,0.0023966678],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9934604,0.0035036015,0.00024028547,0.00032906962,0.002204807,0.00026175287],"domain_scores_gemma":[0.92840046,0.055233143,0.004353084,0.005900087,0.005323053,0.0007900597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007424914,0.00038882912,0.0003527603,0.0014401363,0.00063847756,0.003292751,0.0012162935,0.0009181332,0.0037177538],"category_scores_gemma":[0.045717593,0.00024991212,0.00046665184,0.0013740322,0.0020888478,0.0077310014,0.0023230552,0.0015247558,0.00068462593],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008548188,0.005580978,0.052181523,0.0025042004,0.00036962962,0.0006867093,0.030617377,0.02685591,0.037225027,0.09169578,0.006148459,0.7375862],"study_design_scores_gemma":[0.0014910455,0.02356984,0.1653396,0.003142987,0.0028373348,0.0044177393,0.07253101,0.32064873,0.13300034,0.14732093,0.12505823,0.0006423404],"about_ca_topic_score_codex":0.001969281,"about_ca_topic_score_gemma":0.0021320619,"teacher_disagreement_score":0.007424914,"about_ca_system_score_codex":0.0019321083,"about_ca_system_score_gemma":0.0017120065,"threshold_uncertainty_score":0.039267182},"labels":[],"label_agreement":null},{"id":"W4395464644","doi":"10.1145/3661484","title":"A Meta-Study of Software-Change Intentions","year":2024,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Toronto","funders":"","keywords":"Computer science; Software; Software engineering; Programming language","score_opus":0.3383800416182904,"score_gpt":0.40777684152643945,"score_spread":0.06939679990814906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395464644","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0039083906,0.9876896,0.0020113874,0.0012986388,0.0003173545,0.0005399865,0.0009262905,0.000049897004,0.003258446],"genre_scores_gemma":[0.05879461,0.92776304,0.0065317154,0.0019974262,0.0002791685,0.0017921452,0.0018987495,0.00006981782,0.0008732631],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9814404,0.009678383,0.0038487895,0.0014993425,0.0032367476,0.00029628703],"domain_scores_gemma":[0.7932529,0.170388,0.013931323,0.0059161503,0.015447703,0.0010638995],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.029001852,0.0014279657,0.0035779385,0.026520183,0.00095246447,0.0051663853,0.0018483045,0.0015630763,0.0059630694],"category_scores_gemma":[0.14193238,0.001048754,0.0070141796,0.0221743,0.0013186918,0.006964809,0.0024411296,0.0025269121,0.0010410909],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005650419,0.00015289581,0.0064236317,0.6031352,0.021186193,0.0001664506,0.0037591425,0.00044047413,0.0004480741,0.0070539336,0.0105377,0.34613124],"study_design_scores_gemma":[0.00033946408,0.0005782911,0.015893567,0.7398675,0.09171625,0.00047714816,0.002005407,0.0003003097,0.0008514851,0.007969278,0.13988625,0.00011506005],"about_ca_topic_score_codex":0.005259106,"about_ca_topic_score_gemma":0.014307902,"teacher_disagreement_score":0.97099817,"about_ca_system_score_codex":0.006433797,"about_ca_system_score_gemma":0.013558559,"threshold_uncertainty_score":0.15337825},"labels":[],"label_agreement":null},{"id":"W4395956893","doi":"10.1016/j.infsof.2024.107478","title":"Studying and recommending information highlighting in Stack Overflow answers","year":2024,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Concordia University; Queen's University; University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Disk formatting; Computer science; Information retrieval; Leverage (statistics); Named-entity recognition; Precision and recall; Identification (biology); Source code; Code (set theory); World Wide Web; Artificial intelligence; Data science; Natural language processing; Machine learning; Task (project management); Programming language","score_opus":0.010394210269364759,"score_gpt":0.24737389547338162,"score_spread":0.23697968520401685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395956893","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9402735,0.002736559,0.037122782,0.0022504407,0.00036390164,0.00036054675,0.0021581433,0.0039374777,0.010796543],"genre_scores_gemma":[0.9290985,0.0011826442,0.060050663,0.00036060455,0.00034209405,0.000072992414,0.0028923461,0.00025826678,0.005741953],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99800795,0.00055475475,0.00021863346,0.00025281415,0.0007935332,0.00017232471],"domain_scores_gemma":[0.97104543,0.018641505,0.0035626206,0.0008384344,0.005049034,0.0008629752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022288423,0.00072940986,0.00045920192,0.006686685,0.0010634825,0.0021225733,0.0005839212,0.0014228569,0.0042509725],"category_scores_gemma":[0.042858563,0.00027816033,0.00040786053,0.0027876613,0.00028470898,0.004482593,0.00069288426,0.0008924765,0.0014945207],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015964278,0.0011534278,0.29499704,0.002069005,0.0002379502,0.0010273411,0.0063609583,0.004461471,0.030571984,0.005564542,0.038830426,0.6131295],"study_design_scores_gemma":[0.0002753454,0.0037027386,0.38637868,0.002420893,0.0018551259,0.0024222806,0.019314088,0.27898213,0.10100637,0.02183108,0.18132524,0.00048610324],"about_ca_topic_score_codex":0.0058085485,"about_ca_topic_score_gemma":0.009185128,"teacher_disagreement_score":0.006686685,"about_ca_system_score_codex":0.0006706473,"about_ca_system_score_gemma":0.0015397655,"threshold_uncertainty_score":0.014220893},"labels":[],"label_agreement":null},{"id":"W4396230887","doi":"10.1145/3637307","title":"Characterizing Usability Issue Discussions in Open Source Software Projects","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on Human-Computer Interaction","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Universitas Brawijaya","keywords":"Usability; Computer science; Scope (computer science); Pluralistic walkthrough; Web usability; Context (archaeology); Usability engineering; World Wide Web; Knowledge management; Data science; Human–computer interaction; Geography","score_opus":0.06339403550473677,"score_gpt":0.35596260067153374,"score_spread":0.29256856516679697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396230887","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98638135,0.00037330654,0.008957789,0.00026586666,0.000026332647,0.00023182489,0.00008945884,0.00009032847,0.0035836692],"genre_scores_gemma":[0.9882632,0.00025417787,0.009143562,0.0001048217,0.00003447135,0.0004886715,0.0002787844,0.000064682325,0.0013676016],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9651207,0.020412588,0.0032441574,0.002980627,0.0067627397,0.0014792179],"domain_scores_gemma":[0.75833625,0.17615879,0.035226997,0.0054985653,0.018411346,0.0063680517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02561154,0.0006547094,0.00064594497,0.009072729,0.0036074386,0.0057021347,0.0012143933,0.0021111197,0.001766088],"category_scores_gemma":[0.15400407,0.0005762118,0.0005474968,0.0042430814,0.0021218776,0.0065614698,0.007460089,0.0019037114,0.00049488456],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005066602,0.00070321804,0.35883617,0.0011395735,0.0000975169,0.0011039507,0.48790443,0.000755645,0.009794888,0.0035601805,0.0020230198,0.13357475],"study_design_scores_gemma":[0.000071304,0.001181278,0.49740973,0.0010525243,0.00011341236,0.0016076283,0.44043866,0.011435731,0.005621906,0.010555334,0.030157965,0.00035455343],"about_ca_topic_score_codex":0.0013874684,"about_ca_topic_score_gemma":0.0020263938,"teacher_disagreement_score":0.02561154,"about_ca_system_score_codex":0.0018138421,"about_ca_system_score_gemma":0.0021218199,"threshold_uncertainty_score":0.13544834},"labels":[],"label_agreement":null},{"id":"W4396704595","doi":"10.1007/s00521-024-09855-z","title":"Software effort estimation using convolutional neural network and fuzzy clustering","year":2024,"lang":"en","type":"article","venue":"Neural Computing and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Convolutional neural network; Cluster analysis; Artificial intelligence; Fuzzy clustering; Software; Data mining; Estimation; Fuzzy logic; Artificial neural network; Pattern recognition (psychology); Machine learning; Operating system; Engineering","score_opus":0.02403052001540802,"score_gpt":0.29832974922650823,"score_spread":0.2742992292111002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396704595","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.319087,0.00041112155,0.6771772,0.00010648076,0.000041192034,0.000044326574,0.00022690368,0.00090012315,0.0020056986],"genre_scores_gemma":[0.9354166,0.00007057797,0.062726475,0.000015755282,0.000014442893,0.00002294646,0.00023264723,0.000032546202,0.0014680618],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99952006,0.000058529553,0.000033392353,0.0001631865,0.00015348416,0.00007134096],"domain_scores_gemma":[0.9989579,0.00031410178,0.0001917226,0.00008822297,0.00040398297,0.00004411264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00066808856,0.00068023725,0.0005509904,0.0020866254,0.00030142997,0.00069628906,0.00074695307,0.00074345805,0.0005795131],"category_scores_gemma":[0.0026178255,0.00027581057,0.0005188707,0.0013385186,0.00015425244,0.00084779836,0.00041077787,0.00043468064,0.0002472713],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036049707,0.00031229082,0.040624306,0.00011467766,0.00023590491,0.00013034996,0.00011592471,0.4889183,0.014672018,0.0026051628,0.0020057682,0.44990474],"study_design_scores_gemma":[0.0000017796955,0.000012008655,0.00407499,0.0000035880487,0.000008713243,0.000010520115,0.000007727696,0.9938704,0.0014670985,0.00044542341,0.00009247823,0.0000053566996],"about_ca_topic_score_codex":0.019760117,"about_ca_topic_score_gemma":0.020211024,"teacher_disagreement_score":0.019760117,"about_ca_system_score_codex":0.0010158346,"about_ca_system_score_gemma":0.0006017232,"threshold_uncertainty_score":0.03929025},"labels":[],"label_agreement":null},{"id":"W4396773582","doi":"10.1145/3664606","title":"Unveiling Code Pre-Trained Models: Investigating Syntax and Semantics Capacities","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Research Foundation Singapore","keywords":"Computer science; Syntax; Abstract syntax tree; Abstract syntax; Programming language; Semantics (computer science); Syntax error; Artificial intelligence; Natural language processing; Code (set theory); Source code","score_opus":0.12010431418201437,"score_gpt":0.3292689790807466,"score_spread":0.20916466489873223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396773582","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8486233,0.00048764102,0.14188397,0.00091341103,0.000109787004,0.00012777036,0.0008036361,0.0029508234,0.0040997895],"genre_scores_gemma":[0.9653939,0.00013121555,0.0307589,0.00027426207,0.000015532432,0.00012554083,0.0017018141,0.00025732126,0.0013415],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989538,0.0003099672,0.00005898156,0.0003637743,0.0001836192,0.00012978488],"domain_scores_gemma":[0.98874027,0.0073338975,0.000491947,0.0017437435,0.001314523,0.000375536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002183212,0.0014705306,0.0006558885,0.00078139507,0.00039390047,0.0016733068,0.0018349148,0.0014766719,0.001535535],"category_scores_gemma":[0.021103341,0.0006705432,0.0011003657,0.00056231127,0.001281117,0.0057439497,0.0019641754,0.0041824565,0.0006663609],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000807485,0.00049003697,0.042223092,0.0004541068,0.000353724,0.00044484646,0.0014005534,0.6726412,0.02609305,0.011306734,0.00717723,0.236608],"study_design_scores_gemma":[0.000015829168,0.00009972345,0.0018568499,0.000020053307,0.000041788426,0.00003618636,0.00009231923,0.9869927,0.0057485728,0.004454479,0.00062287407,0.000018634533],"about_ca_topic_score_codex":0.011739079,"about_ca_topic_score_gemma":0.011804596,"teacher_disagreement_score":0.011739079,"about_ca_system_score_codex":0.0017800751,"about_ca_system_score_gemma":0.0018708551,"threshold_uncertainty_score":0.023341477},"labels":[],"label_agreement":null},{"id":"W4396952634","doi":"10.3390/make6020050","title":"Assessment of Software Vulnerability Contributing Factors by Model-Agnostic Explainable AI","year":2024,"lang":"en","type":"article","venue":"Machine Learning and Knowledge Extraction","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Vulnerability (computing); Computer science; Vulnerability assessment; Psychology; Computer security","score_opus":0.013586231344363547,"score_gpt":0.335502881454395,"score_spread":0.32191665011003145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396952634","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27414528,0.0019797136,0.70560503,0.0014201853,0.00005689888,0.00021779953,0.0021941138,0.010071525,0.004309393],"genre_scores_gemma":[0.8171662,0.00043156123,0.1777037,0.00016113873,0.000045374178,0.00013362145,0.0030650874,0.000218504,0.0010748412],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998415,0.00044868444,0.00009813744,0.00037959538,0.0005330267,0.00012556292],"domain_scores_gemma":[0.9928197,0.0042267526,0.001075027,0.00095651636,0.0007921553,0.00012983417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022889639,0.0014811984,0.0006300101,0.0069135586,0.0004084099,0.0013320276,0.0012208702,0.0009808732,0.0012028464],"category_scores_gemma":[0.012420796,0.00031740134,0.0012699891,0.0023093775,0.0006934075,0.0023093598,0.0016805329,0.0014770854,0.00042292575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038334052,0.0005380044,0.14869222,0.0010698581,0.0007513557,0.00046020086,0.00096604053,0.17049345,0.017703673,0.013633948,0.00937629,0.6359315],"study_design_scores_gemma":[0.000021255259,0.00014580014,0.01715273,0.00008170461,0.00018664722,0.0002442089,0.00020003309,0.95082855,0.0065606907,0.02084693,0.003695152,0.000036261732],"about_ca_topic_score_codex":0.0046898467,"about_ca_topic_score_gemma":0.008111654,"teacher_disagreement_score":0.0069135586,"about_ca_system_score_codex":0.0010164377,"about_ca_system_score_gemma":0.0014512858,"threshold_uncertainty_score":0.012105346},"labels":[],"label_agreement":null},{"id":"W4396955972","doi":"10.1017/pds.2024.204","title":"Integrating large language models for improved failure mode and effects analysis (FMEA): a framework and case study","year":2024,"lang":"en","type":"article","venue":"Proceedings of the Design Society","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Rimouski","funders":"","keywords":"Failure mode and effects analysis; Computer science; Reliability engineering; Natural language processing; Engineering","score_opus":0.013323082567451294,"score_gpt":0.29021117517141304,"score_spread":0.27688809260396174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396955972","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18175946,0.00025614162,0.8078155,0.00058154744,0.000021298845,0.0006839015,0.00017677691,0.0031937088,0.0055116704],"genre_scores_gemma":[0.37013373,0.00013331365,0.62781155,0.00005863622,0.0000101534015,0.00027081624,0.00014564599,0.00025366008,0.0011824211],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954104,0.0030446495,0.00016884091,0.00021688825,0.0009715502,0.00018765424],"domain_scores_gemma":[0.98583174,0.011315503,0.00056570064,0.0011495948,0.000982162,0.00015533391],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070612906,0.0010541596,0.0005305884,0.0016942242,0.00089990936,0.0017554985,0.0018543026,0.0015340785,0.0015111545],"category_scores_gemma":[0.011918957,0.0005802562,0.000982973,0.0010897699,0.0015981429,0.00214019,0.0018066241,0.0014308671,0.00026436787],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000820265,0.0032730903,0.012084924,0.0009937462,0.00016527394,0.0032769772,0.006871545,0.523025,0.05616103,0.090394765,0.0042960616,0.2986373],"study_design_scores_gemma":[0.00012111622,0.00058379624,0.0019813774,0.00018890132,0.00010208846,0.00048419277,0.000699417,0.945348,0.024410458,0.011850931,0.0141332755,0.00009641403],"about_ca_topic_score_codex":0.010309045,"about_ca_topic_score_gemma":0.012605663,"teacher_disagreement_score":0.010309045,"about_ca_system_score_codex":0.0017084456,"about_ca_system_score_gemma":0.0017407442,"threshold_uncertainty_score":0.037344098},"labels":[],"label_agreement":null},{"id":"W4397026406","doi":"10.1109/access.2024.3402543","title":"Keeping Deep Learning Models in Check: A History-Based Approach to Mitigate Overfitting","year":2024,"lang":"en","type":"article","venue":"IEEE Access","topic":"Software Engineering Research","field":"Computer Science","cited_by":85,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Huawei Technologies (Canada); University of Alberta","funders":"University of Alberta","keywords":"Overfitting; Computer science; Artificial intelligence; Deep learning; Machine learning; Artificial neural network","score_opus":0.06739076771566944,"score_gpt":0.2979730627291863,"score_spread":0.23058229501351685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4397026406","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1179618,0.0017873249,0.8661626,0.0014941943,0.00022624987,0.0001807041,0.0003197107,0.009049728,0.0028176901],"genre_scores_gemma":[0.8598173,0.0005060358,0.13192348,0.0011353988,0.000174228,0.00018847595,0.00096268376,0.0007185059,0.0045738923],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99731284,0.00049106637,0.0002582952,0.00073280296,0.00090433226,0.00030055235],"domain_scores_gemma":[0.98871464,0.0049603567,0.0017031902,0.002022512,0.0021188592,0.00048051018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053274427,0.002493553,0.0021069502,0.002117478,0.0010790357,0.001943692,0.0038255665,0.0023602052,0.0021144897],"category_scores_gemma":[0.022913251,0.0012377205,0.0015006092,0.0010662794,0.0013075772,0.003921404,0.0028476964,0.0048729964,0.0009874946],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005109867,0.00064162293,0.02613948,0.00024867296,0.00040464086,0.0005785601,0.0005144181,0.47000378,0.011259001,0.003486252,0.008618567,0.47759405],"study_design_scores_gemma":[0.000012770426,0.00014321573,0.0012821227,0.00004356434,0.00006270801,0.00009068304,0.000031252548,0.99017966,0.0042745536,0.0027948585,0.001057833,0.000026792668],"about_ca_topic_score_codex":0.010161386,"about_ca_topic_score_gemma":0.013981239,"teacher_disagreement_score":0.010161386,"about_ca_system_score_codex":0.0014744224,"about_ca_system_score_gemma":0.0024386065,"threshold_uncertainty_score":0.02817458},"labels":[],"label_agreement":null},{"id":"W4397032804","doi":"10.1145/3664597","title":"Deep Learning for Code Intelligence: Survey, Benchmark and Toolkit","year":2024,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Software Engineering Research","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Artificial intelligence; KPI-driven code analysis; Benchmark (surveying); Deep learning; Source code; Code review; Static program analysis; Code (set theory); Machine learning; Software engineering; Data science; Software development; Software; Programming language; Set (abstract data type)","score_opus":0.09497021096444722,"score_gpt":0.38117966236931106,"score_spread":0.28620945140486387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4397032804","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004310011,0.8720589,0.08767967,0.0049299854,0.0008168147,0.00020247723,0.001075528,0.0020945924,0.026831903],"genre_scores_gemma":[0.032301925,0.8863435,0.06463393,0.0020813593,0.0007201888,0.00027893635,0.004388467,0.00071173545,0.008540016],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99813205,0.00032190236,0.00017093432,0.00030175745,0.0009654635,0.00010781865],"domain_scores_gemma":[0.9947463,0.0030766965,0.00025373648,0.00045961735,0.0012924668,0.00017117365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00299255,0.0016407002,0.0011570377,0.0047103995,0.00043474176,0.0020978684,0.002988957,0.0014716492,0.004905615],"category_scores_gemma":[0.011212437,0.0009833323,0.0010311091,0.0058758412,0.0009229009,0.004487081,0.0020069918,0.0031045692,0.0035091883],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000549392,0.00008788248,0.0009376178,0.0055229682,0.00009648569,0.000032247048,0.00007041173,0.008105238,0.0006580967,0.018325953,0.044053774,0.92205435],"study_design_scores_gemma":[0.000047030942,0.00031754456,0.0024288397,0.007672431,0.00025579246,0.00048967614,0.00017150817,0.054202855,0.006399384,0.054477945,0.87341565,0.000121440426],"about_ca_topic_score_codex":0.005694219,"about_ca_topic_score_gemma":0.006200416,"teacher_disagreement_score":0.005694219,"about_ca_system_score_codex":0.0018864747,"about_ca_system_score_gemma":0.0038553653,"threshold_uncertainty_score":0.016410887},"labels":[],"label_agreement":null},{"id":"W4398239158","doi":"10.1145/3639478.3640035","title":"AntiCopyPaster 2.0: Whitebox just-in-time code duplicates extraction","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Plug-in; Computer science; Workflow; Source code; Programming language; Code (set theory); Software engineering; Database; Software; Set (abstract data type)","score_opus":0.021130278101970133,"score_gpt":0.30569740433152387,"score_spread":0.2845671262295537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398239158","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003957966,0.00044410813,0.23617212,0.00030069196,0.00031006258,0.0003765642,0.004557418,0.7494455,0.004435543],"genre_scores_gemma":[0.07077068,0.0010605418,0.55786073,0.0015802297,0.0002500638,0.001297543,0.027221346,0.30669093,0.03326791],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9950996,0.00076937897,0.00048162034,0.0010827346,0.0022799494,0.0002866817],"domain_scores_gemma":[0.98321265,0.007854519,0.0014101334,0.0050254874,0.0019749706,0.00052232086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005680416,0.0036297005,0.0014034684,0.0047829514,0.0010750997,0.0040107714,0.0039331275,0.0023768514,0.02996999],"category_scores_gemma":[0.033077512,0.0027888238,0.0024300076,0.0016603589,0.0011151258,0.0059293406,0.0059471442,0.0029932687,0.025229614],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015300722,0.00034684216,0.0075668553,0.0028310711,0.00036727163,0.0013105099,0.001965106,0.0028315494,0.03967101,0.010057476,0.43481144,0.4967108],"study_design_scores_gemma":[0.00038281863,0.00040742167,0.0071576643,0.0010825297,0.0002581105,0.0027204335,0.0005580735,0.08726683,0.17975736,0.024313625,0.69547135,0.00062374613],"about_ca_topic_score_codex":0.002468395,"about_ca_topic_score_gemma":0.0045986017,"teacher_disagreement_score":0.02996999,"about_ca_system_score_codex":0.000909192,"about_ca_system_score_gemma":0.0025228218,"threshold_uncertainty_score":0.1002596},"labels":[],"label_agreement":null},{"id":"W4398239314","doi":"10.1145/3639478.3643079","title":"A Study of Backporting Code in Open-Source Software for Characterizing Changesets","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Software engineering; Code reuse; Reuse; Open-source software development; Software development; Porting; Personalization; Static program analysis; Software; Slicing; Context (archaeology); Source code; Process (computing); World Wide Web; Programming language; Engineering","score_opus":0.07344489070185992,"score_gpt":0.3514378324231946,"score_spread":0.2779929417213347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398239314","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9867186,0.00013456418,0.01173133,0.000048579725,0.000007333224,0.00006241539,0.0002157305,0.00018712146,0.0008944476],"genre_scores_gemma":[0.97520715,0.00010755303,0.02286539,0.0000240652,0.0000087993485,0.000075832024,0.000936488,0.00018712967,0.00058755436],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9929773,0.0019428582,0.000716982,0.0013298133,0.002634158,0.00039884046],"domain_scores_gemma":[0.9148672,0.05188918,0.0134492265,0.010637642,0.007942577,0.0012141168],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004898723,0.000430616,0.00023477314,0.0043432894,0.0009521401,0.0016295841,0.00090800074,0.0007525743,0.0006837509],"category_scores_gemma":[0.051656686,0.00045545545,0.0004449469,0.0043075383,0.0016033539,0.0036415,0.0018115012,0.0011220482,0.0001822877],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007078388,0.0005779808,0.62428427,0.0010945735,0.00018002761,0.0031046176,0.050088167,0.0135289645,0.06239721,0.009947911,0.0015500867,0.2325383],"study_design_scores_gemma":[0.000029110271,0.001087172,0.8091804,0.00032675243,0.00020999167,0.0033529436,0.014817114,0.09453859,0.049447563,0.008205735,0.018625066,0.00017968581],"about_ca_topic_score_codex":0.0036343548,"about_ca_topic_score_gemma":0.0050084805,"teacher_disagreement_score":0.004898723,"about_ca_system_score_codex":0.00090271543,"about_ca_system_score_gemma":0.00092434813,"threshold_uncertainty_score":0.025907278},"labels":[],"label_agreement":null},{"id":"W4398239316","doi":"10.1145/3639478.3643522","title":"Exploring the Impact of Inheritance on Test Code Maintainability","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Inheritance (genetic algorithm); Computer science; Maintainability; Software engineering; Code reuse; Software evolution; Programming language; Interface (matter); Software development; Software construction; Software; Operating system","score_opus":0.08160541930630628,"score_gpt":0.33081834372151364,"score_spread":0.24921292441520737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398239316","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.994189,0.0003407096,0.0035773923,0.00020829565,0.0000030332803,0.000015879206,0.00020591919,0.00008966633,0.0013700088],"genre_scores_gemma":[0.9954419,0.00013384891,0.0036755188,0.000029654855,0.000005108215,0.000015276632,0.00036712762,0.000048222246,0.0002833333],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9917983,0.0027369137,0.00057378726,0.0008510387,0.0035023256,0.0005375874],"domain_scores_gemma":[0.786923,0.1629269,0.024800278,0.010009684,0.013762076,0.0015780871],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010454075,0.0002625285,0.00021531618,0.0036290006,0.0004951653,0.0014188712,0.00088180526,0.00048165178,0.0007406263],"category_scores_gemma":[0.113047265,0.00027468635,0.0003974957,0.0029397153,0.000968403,0.0028303836,0.0014809113,0.00089378364,0.00016444287],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013968682,0.00012530858,0.9031961,0.00015049311,0.00007529763,0.0003238285,0.0032588611,0.0024756503,0.0033907937,0.0014868061,0.00036528433,0.085011885],"study_design_scores_gemma":[0.000012991214,0.0003263359,0.9639924,0.00013177066,0.00012194932,0.0005072657,0.0015406755,0.021306004,0.0065232143,0.0021869314,0.0033178658,0.000032596246],"about_ca_topic_score_codex":0.009821589,"about_ca_topic_score_gemma":0.013814756,"teacher_disagreement_score":0.010454075,"about_ca_system_score_codex":0.0014127808,"about_ca_system_score_gemma":0.0012265011,"threshold_uncertainty_score":0.055287063},"labels":[],"label_agreement":null},{"id":"W4398239325","doi":"10.1145/3639478.3643126","title":"Recovering Traceability Links between Release Notes and Related Software Artifacts","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Global Institute for Water Security, University of Saskatchewan","keywords":"Traceability; Software engineering; Computer science; Requirements traceability; Transparency (behavior); Software evolution; Software development; Upgrade; Software; Process (computing); Software bug; Software maintenance; Code (set theory); Risk analysis (engineering); Software construction; Computer security; Programming language; Operating system; Business","score_opus":0.02209240958948676,"score_gpt":0.2697056708987591,"score_spread":0.2476132613092723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398239325","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22665918,0.0032455693,0.7155551,0.0021613482,0.0012558703,0.0019229755,0.014877287,0.014441545,0.019881144],"genre_scores_gemma":[0.47635046,0.0019153542,0.47391254,0.0004547264,0.00032207038,0.001065448,0.02845312,0.0023264643,0.015199726],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98151726,0.0031581102,0.0019290152,0.0029476332,0.009873415,0.00057462003],"domain_scores_gemma":[0.85801005,0.046651773,0.02596644,0.038932685,0.028820043,0.001619023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012112213,0.0020830035,0.00076830835,0.011227735,0.0017342989,0.0052837315,0.0024491723,0.0018795351,0.0041805683],"category_scores_gemma":[0.09947615,0.0011767622,0.0006851142,0.007026229,0.00088813115,0.006786939,0.0063137813,0.0032298,0.0029860185],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010410816,0.0010755514,0.07668961,0.002116036,0.00036065647,0.003702702,0.011521086,0.016615234,0.021545274,0.019431565,0.020471025,0.8254302],"study_design_scores_gemma":[0.00039507032,0.0023981968,0.205633,0.004710956,0.0011238771,0.004535168,0.014229782,0.2344443,0.10690811,0.10731035,0.3173216,0.0009895355],"about_ca_topic_score_codex":0.0141646005,"about_ca_topic_score_gemma":0.012881779,"teacher_disagreement_score":0.0141646005,"about_ca_system_score_codex":0.0022419335,"about_ca_system_score_gemma":0.005581069,"threshold_uncertainty_score":0.06405628},"labels":[],"label_agreement":null},{"id":"W4398239408","doi":"10.1145/3639478.3643069","title":"Interpretable Software Maintenance and Support Effort Prediction Using Machine Learning","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Fanshawe College","funders":"","keywords":"Computer science; Machine learning; Software maintenance; Predictive modelling; Software; Support vector machine; Decision tree; Software development; Software engineering; Quality (philosophy); Automation; Artificial intelligence; Engineering","score_opus":0.012734066850915767,"score_gpt":0.25497493416443795,"score_spread":0.24224086731352218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398239408","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3285481,0.000252358,0.66729355,0.00047779238,0.000026736332,0.0000682632,0.0004279223,0.0007309725,0.0021743006],"genre_scores_gemma":[0.9531673,0.00008691168,0.04591896,0.000023452872,0.000011294872,0.000034044253,0.0002876935,0.000016842507,0.000453561],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993818,0.0003002899,0.000040664552,0.00010612279,0.00012832631,0.000042860313],"domain_scores_gemma":[0.9929704,0.005538232,0.00059904356,0.00035254617,0.0004921758,0.000047551177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016620227,0.00055859785,0.00035646785,0.001144204,0.00016758482,0.00092821947,0.00077659776,0.0005076305,0.0011065206],"category_scores_gemma":[0.0095011685,0.00020241602,0.00043470325,0.0009905451,0.00024775698,0.0012374676,0.0003537548,0.0007624526,0.00015977431],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001261719,0.0002841414,0.03542477,0.00012800637,0.00011787508,0.0001849199,0.00034855888,0.8130718,0.0017690563,0.014420508,0.0012422118,0.13288195],"study_design_scores_gemma":[0.0000026403682,0.00001612896,0.0019656501,0.0000067646306,0.0000078866915,0.000011178184,0.000020950843,0.9933257,0.0002856978,0.004233176,0.000119777025,0.000004340857],"about_ca_topic_score_codex":0.0041497773,"about_ca_topic_score_gemma":0.0054270923,"teacher_disagreement_score":0.0041497773,"about_ca_system_score_codex":0.00079878967,"about_ca_system_score_gemma":0.000568077,"threshold_uncertainty_score":0.008789659},"labels":[],"label_agreement":null},{"id":"W4398766359","doi":"10.1145/3639476.3639774","title":"Naturalness of Attention: Revisiting Attention in Code Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Interpretability; Naturalness; Python (programming language); Artificial intelligence; Java; Source code; Programming language; Empirical research; Natural language processing","score_opus":0.020597665641204355,"score_gpt":0.2981587685741326,"score_spread":0.27756110293292824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398766359","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14978477,0.00035410916,0.8413575,0.0023575274,0.000049160815,0.00004608315,0.00012994464,0.00035604273,0.0055648447],"genre_scores_gemma":[0.9501144,0.00023650292,0.047892306,0.00026746592,0.000064844084,0.00008075905,0.00009962259,0.00015040771,0.0010935696],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9984554,0.0008717219,0.000050928644,0.00031283917,0.00020385042,0.000105200634],"domain_scores_gemma":[0.9851795,0.011818432,0.0007214937,0.0013037242,0.00070928916,0.00026767387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003348529,0.00058080535,0.0005058026,0.0010220065,0.00048541004,0.0019104752,0.0015199996,0.0007903158,0.0021651385],"category_scores_gemma":[0.031964287,0.000449481,0.00084549753,0.0005742093,0.00299407,0.0060673803,0.0021044912,0.0026273872,0.00018797783],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019469627,0.00009376716,0.010868145,0.00021396983,0.00011382231,0.00022486516,0.0027642264,0.26135078,0.0067068543,0.6219961,0.0014011277,0.09407157],"study_design_scores_gemma":[0.000007664002,0.000030145837,0.0014202765,0.000026719961,0.000014957325,0.00004165922,0.000105482344,0.62834966,0.0008714864,0.36824438,0.0008753192,0.0000122696965],"about_ca_topic_score_codex":0.0088981865,"about_ca_topic_score_gemma":0.006714337,"teacher_disagreement_score":0.0088981865,"about_ca_system_score_codex":0.0017576389,"about_ca_system_score_gemma":0.001047635,"threshold_uncertainty_score":0.017708957},"labels":[],"label_agreement":null},{"id":"W4398795735","doi":"10.1145/3664646.3665664","title":"AI-Assisted Assessment of Coding Practices in Modern Code Review","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada)","funders":"DeepMind","keywords":"Computer science; Java; Best practice; Software engineering; Python (programming language); Workflow; Coding (social sciences); Programming language; Software deployment; Code review; Static program analysis; Software development; Software; Database","score_opus":0.11175253050054441,"score_gpt":0.4417333888704112,"score_spread":0.3299808583698668,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398795735","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23337372,0.0013407823,0.68193614,0.004134154,0.00059354544,0.00246601,0.001800185,0.058290903,0.016064571],"genre_scores_gemma":[0.415078,0.00034081523,0.57388127,0.0007671839,0.00011075762,0.00075323076,0.0021206406,0.0015187844,0.005429259],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95402724,0.02699235,0.003133522,0.006396483,0.008662686,0.000787628],"domain_scores_gemma":[0.7380848,0.15609053,0.020801688,0.029551981,0.050654534,0.004816477],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.037438326,0.00093446166,0.00096671603,0.0067641903,0.00166864,0.004486866,0.0027986132,0.001621593,0.0038029964],"category_scores_gemma":[0.16906117,0.0009143937,0.00064855,0.002552761,0.0013866428,0.005554151,0.0042454195,0.002968281,0.0031153667],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070805283,0.0010803186,0.045716137,0.0013781643,0.00023154943,0.0004739686,0.011083259,0.022170868,0.029313738,0.0077527254,0.033587996,0.8465032],"study_design_scores_gemma":[0.00050914555,0.00087191886,0.04010433,0.00078468287,0.00015522436,0.000824201,0.004412147,0.7553801,0.058424443,0.04259498,0.09544024,0.0004985912],"about_ca_topic_score_codex":0.004941407,"about_ca_topic_score_gemma":0.014870521,"teacher_disagreement_score":0.96256167,"about_ca_system_score_codex":0.0025865051,"about_ca_system_score_gemma":0.007178938,"threshold_uncertainty_score":0.19799513},"labels":[],"label_agreement":null},{"id":"W4399141822","doi":"10.1109/drcn60692.2024.10539155","title":"BETAC: Bidirectional Encoder Transformer for Assembly Code Function Name Recovery","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; McGill University","funders":"","keywords":"Computer science; Encoder; Transformer; Obfuscation; Binary code; Binary number; Code (set theory); Decoding methods; Resilience (materials science); Distributed computing; Algorithm; Programming language; Computer security; Engineering; Electrical engineering; Arithmetic; Operating system","score_opus":0.023456173892227133,"score_gpt":0.2812545798866184,"score_spread":0.25779840599439124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399141822","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043389592,0.0008414635,0.9083068,0.0005039021,0.00021775639,0.00022069251,0.0013011469,0.040024325,0.0051943725],"genre_scores_gemma":[0.5911823,0.00077679165,0.38607734,0.0004903913,0.00009241196,0.00031278006,0.0042048423,0.0024311033,0.014432136],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99946064,0.000090600624,0.000030155441,0.00014489418,0.00020053047,0.0000731698],"domain_scores_gemma":[0.9980096,0.00069821347,0.00017442217,0.00057995546,0.00046616266,0.000071717805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011276272,0.0011472326,0.0006452518,0.0011286858,0.00047386516,0.00096429075,0.001990993,0.0010884358,0.003348577],"category_scores_gemma":[0.005765372,0.00051114097,0.0008466984,0.0006244878,0.00081838114,0.0025940817,0.00136398,0.0019527393,0.002681043],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008450363,0.000305662,0.008779652,0.00039126578,0.00012072288,0.0006461146,0.0003337006,0.35474804,0.026744263,0.04125512,0.031926177,0.53390425],"study_design_scores_gemma":[0.000012037191,0.000060573024,0.00021926791,0.000014869958,0.00001656227,0.0001543733,0.000015376523,0.9744211,0.014309898,0.0066866786,0.004073277,0.000015902919],"about_ca_topic_score_codex":0.010690233,"about_ca_topic_score_gemma":0.016408442,"teacher_disagreement_score":0.010690233,"about_ca_system_score_codex":0.0012592912,"about_ca_system_score_gemma":0.002772629,"threshold_uncertainty_score":0.02125603},"labels":[],"label_agreement":null},{"id":"W4399204110","doi":"10.1007/978-3-031-55642-5_7","title":"Generative AI for Software Development: A Family of Studies on Code Generation","year":2024,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Generative grammar; Computer science; Code (set theory); Software development; Programming language; Software engineering; Software; Artificial intelligence","score_opus":0.10441114766727508,"score_gpt":0.3304127320367629,"score_spread":0.22600158436948783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399204110","genre_codex":"other","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0105617875,0.05967804,0.2607141,0.0034126136,0.0005825479,0.00011743793,0.00009308081,0.0004895544,0.6643508],"genre_scores_gemma":[0.3379512,0.09110125,0.2107281,0.002881958,0.0012285236,0.000445943,0.00047332392,0.001514228,0.35367548],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99938667,0.00028464795,0.000020368467,0.00008427355,0.00019371949,0.000030327341],"domain_scores_gemma":[0.99587786,0.0036265901,0.000051564395,0.0002782939,0.0001281181,0.00003760588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007525893,0.0007717917,0.00045445605,0.002378884,0.0011469583,0.0033753004,0.0012339966,0.0015295496,0.009217876],"category_scores_gemma":[0.0033906405,0.00050779915,0.00065131724,0.0036368435,0.0051936754,0.0038724374,0.0013479713,0.0022333867,0.0018029081],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008599925,0.000038289625,0.00018164568,0.00030333566,0.0000067005926,0.00006704803,0.002001565,0.001138116,0.00073807186,0.87894773,0.008641642,0.10792721],"study_design_scores_gemma":[0.000012620778,0.00003536777,0.0006754024,0.00050684804,0.000015660655,0.00052680465,0.0010030397,0.007856684,0.00254417,0.6690498,0.3177476,0.00002609473],"about_ca_topic_score_codex":0.0022556002,"about_ca_topic_score_gemma":0.003950464,"teacher_disagreement_score":0.009217876,"about_ca_system_score_codex":0.0017396682,"about_ca_system_score_gemma":0.0012528785,"threshold_uncertainty_score":0.03083688},"labels":[],"label_agreement":null},{"id":"W4399213669","doi":"10.1145/3639477.3639726","title":"Code Impact Beyond Disciplinary Boundaries: Constructing a Multidisciplinary Dependency Graph and Analyzing Cross-Boundary Impact","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ubisoft (Canada); University of Waterloo","funders":"","keywords":"Cross disciplinary; Dependency (UML); Computer science; Multidisciplinary approach; Dependency graph; Graph; Boundary (topology); Data science; Theoretical computer science; Software engineering; Mathematics; Political science","score_opus":0.017570293661934123,"score_gpt":0.36116586944771834,"score_spread":0.3435955757857842,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399213669","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4299719,0.00044284895,0.5544661,0.000792221,0.000042150066,0.00031913587,0.003935675,0.0026476334,0.007382366],"genre_scores_gemma":[0.7185745,0.00035241456,0.272144,0.00010087356,0.000021674796,0.00024400931,0.0064467164,0.000613767,0.001502036],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99883157,0.0003108004,0.000067263834,0.00024096653,0.00045019077,0.0000992004],"domain_scores_gemma":[0.991321,0.0049950653,0.0012728908,0.0009834704,0.0010941213,0.0003334562],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0010752277,0.00068423135,0.00031311242,0.0068784268,0.0009483701,0.0011054897,0.000898366,0.00065710983,0.0009288596],"category_scores_gemma":[0.0093787005,0.0004372792,0.0008009861,0.004712665,0.0010189352,0.0023368474,0.0017394387,0.0009542924,0.00025938888],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001956013,0.00054969813,0.24721602,0.00082693645,0.00031819456,0.0020655715,0.006779918,0.2863043,0.02000551,0.08430915,0.011298739,0.34013033],"study_design_scores_gemma":[0.00003521972,0.00010945688,0.09667575,0.00021576816,0.0002543551,0.0004669646,0.0019215348,0.7411651,0.014831271,0.1137929,0.030430682,0.000100996745],"about_ca_topic_score_codex":0.022805404,"about_ca_topic_score_gemma":0.028203482,"teacher_disagreement_score":0.99312156,"about_ca_system_score_codex":0.0013618189,"about_ca_system_score_gemma":0.0016430715,"threshold_uncertainty_score":0.045345306},"labels":[],"label_agreement":null},{"id":"W4399259577","doi":"10.1007/s10664-024-10464-6","title":"Towards graph-anonymization of software analytics data: empirical study on JIT defect prediction","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Analytics; Data mining; Empirical research; Software; Graph; Data science; Statistics; Theoretical computer science; Mathematics; Programming language","score_opus":0.06589437020513635,"score_gpt":0.3439454061040775,"score_spread":0.27805103589894115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399259577","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70437276,0.00087839214,0.2645387,0.002258876,0.00026286443,0.00048381768,0.019300414,0.0032943212,0.0046098386],"genre_scores_gemma":[0.90428746,0.00038694777,0.07068875,0.0001752918,0.00010550605,0.0001785222,0.022695696,0.00019465364,0.0012872857],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9906985,0.0048358906,0.00048635417,0.0016229973,0.0019236414,0.00043268062],"domain_scores_gemma":[0.9258142,0.03109038,0.00780481,0.029469807,0.0050666747,0.00075417134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052941553,0.00049510226,0.0005464179,0.0045562177,0.0010566104,0.0019705785,0.0012283917,0.0012836949,0.0012443928],"category_scores_gemma":[0.053638384,0.00026905094,0.00066466106,0.00573226,0.001325302,0.0048141163,0.002234711,0.0017562087,0.0008564665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012759748,0.0017373053,0.31383947,0.0012168259,0.0005604559,0.0008581134,0.005413889,0.11060078,0.016827654,0.0606988,0.047514793,0.43945596],"study_design_scores_gemma":[0.000102260274,0.0003318764,0.10062726,0.00032596735,0.00024648264,0.0013425812,0.004372085,0.6858599,0.021262916,0.14234181,0.043061703,0.00012513093],"about_ca_topic_score_codex":0.0036818672,"about_ca_topic_score_gemma":0.0038134577,"teacher_disagreement_score":0.0052941553,"about_ca_system_score_codex":0.00089730433,"about_ca_system_score_gemma":0.0020495018,"threshold_uncertainty_score":0.027998507},"labels":[],"label_agreement":null},{"id":"W4399362314","doi":"10.1007/s10664-024-10487-z","title":"Characterizing and classifying developer forum posts with their intentions","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Mitacs","keywords":"Computer science; Knowledge management","score_opus":0.025739535136037356,"score_gpt":0.2690836191399585,"score_spread":0.24334408400392113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399362314","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9913147,0.00021431618,0.0035382276,0.00013428777,0.00007531432,0.0000620379,0.0010853799,0.00020678488,0.0033689686],"genre_scores_gemma":[0.99105495,0.00013090784,0.0039808718,0.000040213712,0.00011940046,0.00007856152,0.0020575013,0.00006130649,0.002476255],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9972192,0.0007718135,0.00029157868,0.00035885934,0.0010066524,0.0003518907],"domain_scores_gemma":[0.9407309,0.03978971,0.008013896,0.0022183405,0.0066301352,0.0026169477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003196543,0.0005254932,0.0002942241,0.0067015933,0.0007252148,0.0017425946,0.0003581383,0.0009218437,0.0016820993],"category_scores_gemma":[0.03530728,0.0002213058,0.00033247017,0.003117072,0.0003692267,0.0023264594,0.0011114563,0.0007874074,0.001059538],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055288034,0.00043291744,0.8913572,0.00031687148,0.00009096902,0.00028412763,0.0029757188,0.00062797393,0.016032238,0.0011200089,0.004858622,0.081350386],"study_design_scores_gemma":[0.000035236535,0.00064891647,0.9508951,0.00012271707,0.00014478633,0.0006607403,0.005017277,0.024967857,0.006279759,0.0019326214,0.0092207985,0.00007423537],"about_ca_topic_score_codex":0.0013656802,"about_ca_topic_score_gemma":0.0037126583,"teacher_disagreement_score":0.0067015933,"about_ca_system_score_codex":0.00031044116,"about_ca_system_score_gemma":0.0005903927,"threshold_uncertainty_score":0.016905129},"labels":[],"label_agreement":null},{"id":"W4399530762","doi":"10.1145/3644815.3644965","title":"DVC in Open Source ML-development: The Action and the Reaction","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Calgary; Concordia University","funders":"","keywords":"Software versioning; Computer science; Popularity; Software engineering; Process (computing); Quality (philosophy); Software; Open-source software development; Software development; Database; World Wide Web; Operating system","score_opus":0.03664933410204949,"score_gpt":0.3072741403339433,"score_spread":0.2706248062318938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399530762","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.87678134,0.003994267,0.056181986,0.023428446,0.0004021226,0.0004900842,0.00030488626,0.002774589,0.03564234],"genre_scores_gemma":[0.974698,0.00086661207,0.017160099,0.0026311935,0.00022394682,0.00017296906,0.00023913816,0.0008512539,0.0031568],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.93441784,0.02811526,0.0033306284,0.00905759,0.021698063,0.0033806884],"domain_scores_gemma":[0.66846186,0.20251703,0.049088247,0.036403153,0.031020613,0.012509148],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.039056186,0.0010176455,0.00060980953,0.004991284,0.0027752155,0.009909337,0.0032681355,0.0036678414,0.0031971007],"category_scores_gemma":[0.23713975,0.0012667611,0.0007585038,0.004316352,0.007551285,0.011886985,0.0093822405,0.005925174,0.0012344439],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045330226,0.0007329297,0.41053963,0.0010848513,0.00017071128,0.0019366744,0.06100424,0.0044684876,0.013016314,0.03831658,0.016933056,0.45134327],"study_design_scores_gemma":[0.00019900645,0.0017082787,0.55436987,0.0024168645,0.00023777783,0.005656672,0.06351774,0.045011465,0.017084628,0.070236646,0.23873614,0.0008249076],"about_ca_topic_score_codex":0.005861086,"about_ca_topic_score_gemma":0.0044728406,"teacher_disagreement_score":0.9967319,"about_ca_system_score_codex":0.0066419905,"about_ca_system_score_gemma":0.0048912717,"threshold_uncertainty_score":0.20655131},"labels":[],"label_agreement":null},{"id":"W4399567378","doi":"10.1145/3650105.3652299","title":"Fine Tuning Large Language Model for Secure Code Generation","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Concordia University","funders":"","keywords":"Computer science; Code (set theory); Code generation; Vulnerability (computing); Face (sociological concept); Source code; Programming language; Computer security","score_opus":0.033537094149657984,"score_gpt":0.312070987112272,"score_spread":0.278533892962614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399567378","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34289935,0.0036595832,0.47699708,0.0023166803,0.0011915063,0.00075306016,0.0055824225,0.15599129,0.01060904],"genre_scores_gemma":[0.63903975,0.0006772025,0.32702494,0.002046398,0.00012594368,0.00073991634,0.01645758,0.007024407,0.006863843],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985341,0.00041783837,0.000096250675,0.0005093142,0.00026708038,0.00017530532],"domain_scores_gemma":[0.9962115,0.0018020084,0.00020368947,0.0008979088,0.0007125066,0.00017232666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001942916,0.0016413477,0.0008029478,0.0010707021,0.0005170935,0.0014398593,0.0025583534,0.0014075281,0.0036458136],"category_scores_gemma":[0.010661581,0.0007886994,0.0016393282,0.0007348386,0.0007054153,0.0031833004,0.001584016,0.0035103224,0.003353844],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080011,0.0009831347,0.015109117,0.00084346737,0.0003517591,0.0007073968,0.00053952576,0.46763772,0.04772319,0.0052928925,0.067612596,0.39239913],"study_design_scores_gemma":[0.00010391215,0.00011722331,0.00070786953,0.000035284906,0.00006138721,0.00014019386,0.00007180561,0.9719416,0.014283167,0.0038938902,0.008606552,0.00003714874],"about_ca_topic_score_codex":0.008809023,"about_ca_topic_score_gemma":0.016157996,"teacher_disagreement_score":0.008809023,"about_ca_system_score_codex":0.0014652434,"about_ca_system_score_gemma":0.0028050654,"threshold_uncertainty_score":0.01751554},"labels":[],"label_agreement":null},{"id":"W4399572687","doi":"10.1145/3641822.3641873","title":"Charting a Path to Efficient Onboarding: The Role of Software Visualization","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Onboarding; Computer science; Visualization; Software development; Software engineering; Process (computing); Software; Knowledge management; Software visualization; Personal software process; Context (archaeology); Software inspection; Software construction; Data science; Software quality; Artificial intelligence","score_opus":0.01078796831382044,"score_gpt":0.27624611249139436,"score_spread":0.2654581441775739,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399572687","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18631306,0.018078119,0.57248735,0.090581454,0.0017400075,0.00057625334,0.00016847323,0.0032579796,0.12679724],"genre_scores_gemma":[0.7237305,0.0109459115,0.25467673,0.0019793161,0.00043839173,0.00027914337,0.00012668717,0.0005503947,0.0072729057],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9820441,0.013058478,0.00052812044,0.0012916697,0.002271148,0.0008065376],"domain_scores_gemma":[0.93603444,0.042109534,0.005446682,0.0048813643,0.007139401,0.004388607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018173348,0.0011510821,0.00043347632,0.0026769992,0.003613336,0.014249339,0.002262457,0.0027817332,0.008107953],"category_scores_gemma":[0.040634554,0.00066355977,0.0005680304,0.0023340331,0.01026679,0.014805857,0.0057766745,0.00429463,0.0015716979],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001725408,0.00037832066,0.015137233,0.003085707,0.000057894158,0.0008916285,0.07849887,0.0030592189,0.009626585,0.17494407,0.030672643,0.6834754],"study_design_scores_gemma":[0.00008053804,0.0010687435,0.0213538,0.0067155794,0.00012676009,0.002519547,0.09267243,0.015405164,0.011914908,0.2854516,0.5622908,0.00040013646],"about_ca_topic_score_codex":0.0029038219,"about_ca_topic_score_gemma":0.002043953,"teacher_disagreement_score":0.018173348,"about_ca_system_score_codex":0.0028872713,"about_ca_system_score_gemma":0.0072876005,"threshold_uncertainty_score":0.096111},"labels":[],"label_agreement":null},{"id":"W4399577198","doi":"10.1145/3650105.3652290","title":"PathOCL: Path-Based Prompt Augmentation for OCL Generation with GPT-4","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Unified Modeling Language; Computer science; Object Constraint Language; Programming language; Applications of UML; UML tool; Context (archaeology); Software engineering; Class (philosophy); Artificial intelligence; Software","score_opus":0.0277043368443298,"score_gpt":0.2857301926395388,"score_spread":0.258025855795209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399577198","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0092681255,0.00012028122,0.88092756,0.00013173156,0.00013460057,0.00046912252,0.0008466744,0.10545968,0.0026422075],"genre_scores_gemma":[0.0978395,0.00013341474,0.8834676,0.00017123834,0.00003494493,0.0006232165,0.0023884831,0.0107496185,0.004591911],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99847955,0.00039374758,0.00013523888,0.0003470669,0.00053046766,0.00011395317],"domain_scores_gemma":[0.9940965,0.0031170833,0.00039436275,0.0013119426,0.00091038423,0.00016974908],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015891107,0.0016711456,0.00051961927,0.0011397917,0.00043089132,0.0011152985,0.002184388,0.0011004395,0.02540776],"category_scores_gemma":[0.011465398,0.0008406747,0.00082929444,0.00064278534,0.00089713116,0.002269229,0.0031049794,0.0015057893,0.0070735444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015906964,0.00029607074,0.00303203,0.0015717617,0.00007995919,0.0012415132,0.002731974,0.020402392,0.102535166,0.01547113,0.05152922,0.79951805],"study_design_scores_gemma":[0.0005167389,0.0009352044,0.0017251716,0.00022766943,0.00013564329,0.001732924,0.0006502725,0.55836827,0.21305606,0.023195764,0.19924459,0.00021176122],"about_ca_topic_score_codex":0.001699939,"about_ca_topic_score_gemma":0.0022022706,"teacher_disagreement_score":0.02540776,"about_ca_system_score_codex":0.00061913306,"about_ca_system_score_gemma":0.0014686813,"threshold_uncertainty_score":0.084997416},"labels":[],"label_agreement":null},{"id":"W4399631545","doi":"10.1145/3643916.3644439","title":"TerraMetrics: An Open Source Tool for Infrastructure-as-Code (IaC) Quality Metrics in Terraform","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"DevOps; Computer science; Source code; Software engineering; JSON; Software; Quality assurance; Software quality; JavaScript; Java; Abstract syntax tree; Interpreter; Quality (philosophy); Metric (unit); Software quality assurance; Software deployment; Parsing; World Wide Web; Operating system; Programming language; Software development; Service (business); Engineering","score_opus":0.056479815997341185,"score_gpt":0.38034339799072053,"score_spread":0.32386358199337933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399631545","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012278068,0.00054166885,0.31615132,0.00055172347,0.00017740857,0.0009098225,0.036485516,0.6160206,0.016883792],"genre_scores_gemma":[0.108115174,0.0011786458,0.5394449,0.0006852361,0.00013260657,0.0022206374,0.1524343,0.17790776,0.017880714],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99599755,0.00066347956,0.0006579347,0.00049675524,0.0019693882,0.00021497403],"domain_scores_gemma":[0.9884371,0.00441002,0.001848541,0.0022931974,0.0026845643,0.00032661858],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004469114,0.0019632175,0.00072952243,0.0056706853,0.00079432386,0.0031643256,0.0017796592,0.0009837642,0.017198246],"category_scores_gemma":[0.025942788,0.0013185116,0.001464974,0.004043754,0.0009182526,0.0036495603,0.0029201102,0.0017631642,0.009070513],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050902006,0.00024342847,0.015868105,0.0028327247,0.00023116068,0.0007570115,0.0026662108,0.01722606,0.016054371,0.027924202,0.4160384,0.49964917],"study_design_scores_gemma":[0.00024301329,0.00031450955,0.02067426,0.0018108444,0.00012559949,0.0013949551,0.00087347184,0.10017323,0.036823954,0.032987818,0.80402696,0.00055134436],"about_ca_topic_score_codex":0.011056988,"about_ca_topic_score_gemma":0.012693052,"teacher_disagreement_score":0.9955309,"about_ca_system_score_codex":0.0014277195,"about_ca_system_score_gemma":0.0032598672,"threshold_uncertainty_score":0.0575338},"labels":[],"label_agreement":null},{"id":"W4399644038","doi":"10.1007/s10664-024-10457-5","title":"Utilization of pre-trained language models for adapter-based knowledge transfer in software engineering","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"","keywords":"Adapter (computing); Computer science; Code (set theory); Automatic summarization; Downstream (manufacturing); Software; Artificial intelligence; Natural language processing; Source code; Programming language; Engineering; Computer hardware","score_opus":0.046160511992757766,"score_gpt":0.31951923156039336,"score_spread":0.2733587195676356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399644038","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32993022,0.00051434815,0.637453,0.0005396535,0.00032052258,0.00062240934,0.0009771938,0.018688751,0.010953998],"genre_scores_gemma":[0.8385911,0.00025424233,0.1538385,0.00029415224,0.00004879279,0.0004925391,0.0021710827,0.000576992,0.0037325718],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980392,0.00081568153,0.00016980113,0.0005940477,0.00023328981,0.00014806434],"domain_scores_gemma":[0.985067,0.0092957895,0.00041206923,0.002288898,0.0026167247,0.00031955453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037321278,0.001343607,0.0007218871,0.0014722736,0.0006353962,0.0020035126,0.0021061231,0.0014426516,0.004733868],"category_scores_gemma":[0.024499623,0.0006136435,0.0009057356,0.0010398434,0.0005029205,0.006184866,0.0034896436,0.002789681,0.0031158703],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008343216,0.0013152974,0.010721737,0.00041529493,0.00019983425,0.00045292376,0.0016617961,0.0588063,0.029607778,0.0034365996,0.006341578,0.88620645],"study_design_scores_gemma":[0.000083711624,0.0004824591,0.003448649,0.00010598523,0.000198966,0.0002991454,0.00065086153,0.9351079,0.04640855,0.0091023715,0.004035414,0.00007592981],"about_ca_topic_score_codex":0.00638869,"about_ca_topic_score_gemma":0.0067452933,"teacher_disagreement_score":0.00638869,"about_ca_system_score_codex":0.0010652762,"about_ca_system_score_gemma":0.00225697,"threshold_uncertainty_score":0.019737601},"labels":[],"label_agreement":null},{"id":"W4399668026","doi":"10.1145/3661167.3661234","title":"How Much Logs Does My Source Code File Need? Learning to Predict the Density of Logs","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Computer science; Source code; Code (set theory); Programming language; Database; Set (abstract data type)","score_opus":0.011591529814095612,"score_gpt":0.2375850552010579,"score_spread":0.2259935253869623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399668026","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91359496,0.0004879734,0.07467375,0.0010613122,0.000046516652,0.00009405812,0.0052297665,0.0029570127,0.0018547587],"genre_scores_gemma":[0.97645265,0.00018215987,0.016927512,0.00007442626,0.000030851108,0.000069406255,0.0052952706,0.00006278662,0.0009048973],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991579,0.00016750641,0.00006382363,0.0003425689,0.00017030886,0.00009786844],"domain_scores_gemma":[0.99085325,0.005529344,0.0015502018,0.0005413959,0.0011393654,0.00038643694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012000832,0.0010134474,0.000498092,0.0025786024,0.0003326499,0.0011131798,0.0007109942,0.00078294653,0.00105478],"category_scores_gemma":[0.013300476,0.00037222388,0.00043346346,0.0013616891,0.00041756447,0.002843126,0.0007878791,0.0013707171,0.0011281912],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003487423,0.00051924284,0.7242229,0.00025348202,0.000095616,0.00022877884,0.0006440354,0.0743481,0.0029042063,0.0009378375,0.011041062,0.18445595],"study_design_scores_gemma":[0.000012414863,0.00009843962,0.10705796,0.000048105932,0.00002833239,0.0001727076,0.00038881137,0.88500124,0.0026055474,0.0027844147,0.0017675579,0.00003437432],"about_ca_topic_score_codex":0.010183626,"about_ca_topic_score_gemma":0.0171378,"teacher_disagreement_score":0.010183626,"about_ca_system_score_codex":0.0006846367,"about_ca_system_score_gemma":0.0007507814,"threshold_uncertainty_score":0.020248711},"labels":[],"label_agreement":null},{"id":"W4399668128","doi":"10.1145/3661167.3661190","title":"Exploring Influence of Feature Toggles on Code Complexity","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Feature (linguistics); Code (set theory); Programming language","score_opus":0.15958671814062309,"score_gpt":0.3188125642577183,"score_spread":0.15922584611709523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399668128","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9971102,0.00007752492,0.002037944,0.000030583153,0.0000031629788,0.000018154635,0.00014424097,0.00015631948,0.0004218023],"genre_scores_gemma":[0.9966421,0.000046598292,0.0026061416,0.000010175494,0.0000036452896,0.000027006448,0.00032661806,0.000072978975,0.0002647181],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9955771,0.0012312807,0.0002791765,0.0010231861,0.0015707364,0.00031860766],"domain_scores_gemma":[0.85954225,0.11031037,0.015226958,0.005761589,0.0074217883,0.0017371003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026400303,0.00059056043,0.00033365423,0.0027847202,0.00035406824,0.0011491708,0.00051277404,0.0004147862,0.0007180862],"category_scores_gemma":[0.05016388,0.00032522058,0.000508494,0.0020111364,0.00067245343,0.001631305,0.0008417611,0.0007462196,0.00019536448],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006212486,0.00027430136,0.90382326,0.00021147929,0.0002621634,0.00045937192,0.0017102897,0.012191459,0.019575102,0.0003865229,0.0006113917,0.059873365],"study_design_scores_gemma":[0.000018011922,0.0006847138,0.92862177,0.000031027925,0.00016016781,0.00034641515,0.0008561634,0.056615733,0.011036503,0.00060699374,0.0009651244,0.000057379013],"about_ca_topic_score_codex":0.0037703323,"about_ca_topic_score_gemma":0.0052776677,"teacher_disagreement_score":0.0037703323,"about_ca_system_score_codex":0.00049273437,"about_ca_system_score_gemma":0.00067740044,"threshold_uncertainty_score":0.013961971},"labels":[],"label_agreement":null},{"id":"W4399686980","doi":"10.1007/s10664-024-10500-5","title":"Common challenges of deep reinforcement learning applications development: an empirical study","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Fonds de recherche du Québec; Canadian Institute for Advanced Research","keywords":"Computer science; Taxonomy (biology); Leverage (statistics); Popularity; Reinforcement learning; Artificial intelligence; Data science; Software engineering","score_opus":0.039303453045082565,"score_gpt":0.3277676583422513,"score_spread":0.28846420529716876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399686980","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.987041,0.00031186646,0.0067943605,0.0009836364,0.000018697485,0.0001066572,0.00008206598,0.000091182636,0.00457047],"genre_scores_gemma":[0.9960871,0.00010104471,0.002775226,0.0000744711,0.000007262242,0.00004046179,0.00008242083,0.000022175567,0.00080985343],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9903364,0.0044991337,0.0006255725,0.0010198565,0.0029899962,0.0005290202],"domain_scores_gemma":[0.82047254,0.13490298,0.013224609,0.012646364,0.014958517,0.0037949746],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010790749,0.0003709958,0.00035440119,0.0010528815,0.0011102257,0.002402581,0.0013159121,0.0015164483,0.002389452],"category_scores_gemma":[0.12243964,0.0003613032,0.00028923724,0.0012688949,0.0014179333,0.004202162,0.0022460246,0.0023739596,0.00055567373],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018156403,0.006771544,0.48493907,0.00096737366,0.00018547906,0.0012108996,0.015945604,0.030475266,0.0074069574,0.02877593,0.010908402,0.41059783],"study_design_scores_gemma":[0.00036116663,0.004829712,0.3836811,0.0009696836,0.00024451723,0.0025596744,0.030809632,0.4510855,0.016185759,0.06619454,0.042821106,0.00025769163],"about_ca_topic_score_codex":0.0026124239,"about_ca_topic_score_gemma":0.0035328737,"teacher_disagreement_score":0.98920923,"about_ca_system_score_codex":0.0017315651,"about_ca_system_score_gemma":0.0025650982,"threshold_uncertainty_score":0.057067633},"labels":[],"label_agreement":null},{"id":"W4399860893","doi":"10.1145/3660650.3660673","title":"Can You Spot the AI? Incorporating GenAI into Technical Writing Assignments","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Reflection (computer programming); Technical writing; Writing process; Process (computing); Mathematics education; Artificial intelligence; Psychology; Higher education; Programming language","score_opus":0.017394745913907196,"score_gpt":0.290937253681044,"score_spread":0.2735425077671368,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399860893","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7130005,0.00049539964,0.22572015,0.0046161814,0.00075191195,0.0021293901,0.0002721173,0.009703252,0.04331116],"genre_scores_gemma":[0.6366568,0.0003664653,0.34434494,0.0012139074,0.00015593994,0.0013939525,0.0003455728,0.0005720972,0.014950274],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9901245,0.0060064415,0.00053059513,0.0012826218,0.001577633,0.00047826234],"domain_scores_gemma":[0.92320275,0.056125294,0.005465586,0.0071730986,0.0052578426,0.0027753797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011900289,0.0008073982,0.0005955373,0.0010442267,0.0009526945,0.0035579232,0.0016301468,0.0010120445,0.0065285987],"category_scores_gemma":[0.07412594,0.00045771062,0.00031164542,0.00063065835,0.0011804636,0.003399422,0.0040432876,0.0018381936,0.0031720179],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010730617,0.0022835312,0.02422778,0.0012118911,0.000045831865,0.0006609366,0.046160802,0.0023681682,0.042811487,0.006922224,0.01536709,0.85686725],"study_design_scores_gemma":[0.0009628463,0.014182717,0.106473304,0.0030493299,0.00032370468,0.005137512,0.06906056,0.07391101,0.13713913,0.09347495,0.49530533,0.0009795954],"about_ca_topic_score_codex":0.00021484103,"about_ca_topic_score_gemma":0.0009437537,"teacher_disagreement_score":0.011900289,"about_ca_system_score_codex":0.00069736614,"about_ca_system_score_gemma":0.0015656352,"threshold_uncertainty_score":0.06293553},"labels":[],"label_agreement":null},{"id":"W4400065154","doi":"10.1145/3643787.3648028","title":"Aligning Programming Language and Natural Language: Exploring Design Choices in Multi-Modal Transformer-Based Embedding for Bug Localization","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Embedding; Computer science; Source code; Artificial intelligence; Natural language; Software; Machine learning; Natural language processing; Programming language","score_opus":0.061910445171532276,"score_gpt":0.34671907705832894,"score_spread":0.28480863188679667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400065154","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15843682,0.0006005728,0.837885,0.0005349956,0.000028140455,0.000056067187,0.00012803076,0.001037106,0.0012933215],"genre_scores_gemma":[0.8477858,0.0002267507,0.150013,0.00014389679,0.000014310599,0.000083030085,0.00033149854,0.00024393568,0.001157682],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99890614,0.00057356467,0.000050801802,0.00024947,0.00012946814,0.00009051031],"domain_scores_gemma":[0.9960443,0.0028046018,0.00031597665,0.00038103297,0.00033836943,0.000115786796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001829582,0.0007236976,0.00047148592,0.000742585,0.0002604525,0.0010212965,0.0010687322,0.000704601,0.0018495101],"category_scores_gemma":[0.009827462,0.00040471234,0.00061042473,0.0005641693,0.0013029247,0.0045844144,0.0017611244,0.0013985337,0.00047178808],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011632167,0.00061612605,0.01407255,0.00074377906,0.0001753143,0.00047856945,0.0021176003,0.3398573,0.03782449,0.066200785,0.004224285,0.53252596],"study_design_scores_gemma":[0.000016629112,0.00009748189,0.00050589413,0.000026083539,0.000023879922,0.00007315742,0.00017217566,0.95515245,0.0032609683,0.039965294,0.00069246325,0.000013489107],"about_ca_topic_score_codex":0.0018701811,"about_ca_topic_score_gemma":0.0026958995,"teacher_disagreement_score":0.0018701811,"about_ca_system_score_codex":0.00079602154,"about_ca_system_score_gemma":0.0006210062,"threshold_uncertainty_score":0.00967586},"labels":[],"label_agreement":null},{"id":"W4400242114","doi":"10.1145/3643991.3644920","title":"Enhancing Performance Bug Prediction Using Performance Code Metrics","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Toronto; IBM (Canada)","funders":"","keywords":"Computer science; Software bug; Code (set theory); Programming language; Performance prediction; Software","score_opus":0.029421121564483114,"score_gpt":0.27464007787532047,"score_spread":0.24521895631083734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400242114","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.720529,0.0031887891,0.24324527,0.0014363145,0.0002284906,0.00024005515,0.0040988037,0.02250933,0.0045238533],"genre_scores_gemma":[0.94201916,0.00042777066,0.051006872,0.00017571996,0.00006997286,0.000087270615,0.004568582,0.00040587044,0.001238879],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970788,0.0005362195,0.00023331575,0.0005877495,0.0012356924,0.00032830294],"domain_scores_gemma":[0.9781502,0.0082438635,0.0048935236,0.0016483543,0.0062520397,0.00081209274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034956483,0.002192998,0.0010234285,0.005580915,0.00033462935,0.0014871965,0.00094280014,0.0009601516,0.00065748097],"category_scores_gemma":[0.02342481,0.00047000652,0.00079702836,0.0025958675,0.00039728277,0.0023913158,0.0009590852,0.0014495227,0.0010296229],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002481861,0.00062042323,0.5921922,0.00036474154,0.0001935461,0.00035465084,0.0002672191,0.1064523,0.006943598,0.0012222405,0.013699456,0.2774415],"study_design_scores_gemma":[0.00001652149,0.00019967824,0.061118893,0.00007121662,0.000059826798,0.00019333446,0.00007501765,0.9272098,0.006562579,0.001587104,0.0028531333,0.000052927433],"about_ca_topic_score_codex":0.008813991,"about_ca_topic_score_gemma":0.00878902,"teacher_disagreement_score":0.008813991,"about_ca_system_score_codex":0.0006436669,"about_ca_system_score_gemma":0.0015038886,"threshold_uncertainty_score":0.018486977},"labels":[],"label_agreement":null},{"id":"W4400242190","doi":"10.1145/3643991.3644881","title":"Multi-faceted Code Smell Detection at Scale using DesigniteJava 2.0","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Scale (ratio); Code (set theory); Artificial intelligence; Programming language; Geography; Cartography","score_opus":0.05527756510124146,"score_gpt":0.30714638548599854,"score_spread":0.2518688203847571,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400242190","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04976556,0.0006684073,0.2873675,0.00061435613,0.00019476561,0.00045865987,0.010733098,0.6440864,0.0061112703],"genre_scores_gemma":[0.34672776,0.00061875273,0.51042414,0.00077573536,0.00008239322,0.001701721,0.039693944,0.08959846,0.010377195],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99576086,0.000860759,0.00046864513,0.0010807416,0.001570386,0.00025866766],"domain_scores_gemma":[0.9801799,0.009704872,0.0022963446,0.0048849424,0.0024590679,0.0004748257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050717727,0.002206301,0.00087055826,0.0047062165,0.0007409609,0.0032200757,0.0023806954,0.0011677331,0.0049903844],"category_scores_gemma":[0.021651404,0.0024544005,0.0020939908,0.0018371616,0.0011076644,0.0039011952,0.0037532307,0.003020895,0.0028042758],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025927182,0.00094393117,0.12119104,0.0044377167,0.0013238427,0.0018148825,0.0056643,0.03315108,0.080330685,0.02029366,0.24143599,0.48682025],"study_design_scores_gemma":[0.00053109776,0.0008571135,0.044950657,0.00087666314,0.00041102467,0.002157084,0.00093381654,0.59792286,0.112212814,0.03831951,0.20016353,0.0006638149],"about_ca_topic_score_codex":0.003039018,"about_ca_topic_score_gemma":0.00616547,"teacher_disagreement_score":0.0050717727,"about_ca_system_score_codex":0.0009766198,"about_ca_system_score_gemma":0.0022884782,"threshold_uncertainty_score":0.026822448},"labels":[],"label_agreement":null},{"id":"W4400242886","doi":"10.1145/3643991.3644864","title":"CodeLL: A Lifelong Learning Dataset to Support the Co-Evolution of Data and Language Models of Code","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Lifelong learning; Code (set theory); Data modeling; Artificial intelligence; Natural language processing; Data science; Programming language; Software engineering; Psychology","score_opus":0.05609616902081996,"score_gpt":0.3477873227265498,"score_spread":0.29169115370572984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400242886","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13426428,0.0020783842,0.03163119,0.0015173573,0.00040736428,0.0005674267,0.8039822,0.01816491,0.0073869172],"genre_scores_gemma":[0.0631723,0.00034709115,0.034177214,0.00029169224,0.00006922598,0.00082531176,0.89718556,0.0009901886,0.002941428],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99759454,0.0004291187,0.00031203986,0.00087329745,0.00059968734,0.00019126074],"domain_scores_gemma":[0.99122417,0.0030592487,0.0011361778,0.002126396,0.0017475932,0.00070648146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017034829,0.001186523,0.00061748986,0.0043661217,0.0009319041,0.0015798714,0.002056108,0.0019973894,0.0031797907],"category_scores_gemma":[0.013677681,0.00041748228,0.001017629,0.0040928056,0.0006552381,0.0022819582,0.0024139062,0.0024652644,0.005254027],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084463856,0.0015749886,0.14931405,0.003025315,0.00040504214,0.00081028935,0.002014055,0.016136259,0.010651465,0.005670737,0.58234286,0.22721028],"study_design_scores_gemma":[0.00036157185,0.00058423454,0.13046241,0.00051296887,0.00013940249,0.0010781445,0.001177199,0.09258204,0.013522969,0.009456989,0.74983096,0.00029110664],"about_ca_topic_score_codex":0.0113544855,"about_ca_topic_score_gemma":0.03154083,"teacher_disagreement_score":0.0113544855,"about_ca_system_score_codex":0.0011587319,"about_ca_system_score_gemma":0.0018802616,"threshold_uncertainty_score":0.02257675},"labels":[],"label_agreement":null},{"id":"W4400266869","doi":"10.1145/3643991.3644907","title":"PeaTMOSS: A Dataset and Initial Analysis of Pre-Trained Models in Open-Source Software","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Science Foundation","keywords":"Computer science; Open source software; Open source; Software; Artificial intelligence; Programming language","score_opus":0.035562403760425194,"score_gpt":0.3328252172419072,"score_spread":0.29726281348148204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400266869","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34725615,0.0037243564,0.0379922,0.002482716,0.001071678,0.00047350893,0.54086536,0.05053747,0.015596668],"genre_scores_gemma":[0.18819863,0.000957053,0.03229921,0.00046951312,0.00011222938,0.0006092572,0.76957434,0.0024761988,0.0053036143],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981964,0.00035468597,0.00017900145,0.0004721762,0.00060064974,0.00019710885],"domain_scores_gemma":[0.9939619,0.0024886574,0.0003762971,0.001604211,0.0012211226,0.00034777622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020955398,0.0017910337,0.0006373668,0.0029376685,0.00087815995,0.0012197916,0.002426606,0.0022066275,0.0039425557],"category_scores_gemma":[0.01236333,0.00058493443,0.0016204377,0.0022853415,0.0010095543,0.0019642725,0.002141502,0.00254864,0.0051466427],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016061204,0.0013795869,0.06892976,0.002770936,0.0007159352,0.0019751927,0.0009971659,0.10194959,0.0092711,0.0068222666,0.635078,0.16850446],"study_design_scores_gemma":[0.0008084566,0.0011899397,0.11680407,0.0010305109,0.00032871758,0.001541742,0.001426495,0.39555463,0.02519174,0.024137178,0.43157536,0.00041112668],"about_ca_topic_score_codex":0.025397753,"about_ca_topic_score_gemma":0.063849196,"teacher_disagreement_score":0.025397753,"about_ca_system_score_codex":0.0015612771,"about_ca_system_score_gemma":0.0020321482,"threshold_uncertainty_score":0.050499856},"labels":[],"label_agreement":null},{"id":"W4400289352","doi":"10.1007/s10664-024-10476-2","title":"Design smells in multi-language systems and bug-proneness: a survival analysis","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Natural language processing; Programming language; Linguistics; Artificial intelligence","score_opus":0.050847855513788974,"score_gpt":0.3200556375862844,"score_spread":0.26920778207249546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400289352","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99750525,0.00010983171,0.0018684545,0.0000458291,0.0000020335694,0.000006017773,0.000085185755,0.000038190006,0.00033919633],"genre_scores_gemma":[0.9992454,0.000023714108,0.00042479992,0.0000049325677,0.0000022825293,0.0000066778066,0.00007396958,0.000012009797,0.00020636791],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99825436,0.0006952458,0.0001424062,0.00028809457,0.00041620826,0.00020356034],"domain_scores_gemma":[0.89172286,0.0746526,0.020106262,0.00537378,0.006193399,0.0019511753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006945815,0.00045381923,0.0004444632,0.0034781678,0.00056042866,0.001384672,0.00065695425,0.0008309903,0.0032929387],"category_scores_gemma":[0.04203862,0.00032597428,0.0015222736,0.002380495,0.0010659051,0.0023733093,0.0012226589,0.0014094951,0.0004726969],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003155752,0.00015771456,0.9744241,0.000052556403,0.00023338127,0.00022722031,0.001234919,0.0033241722,0.0017032181,0.00079228374,0.000180681,0.017354203],"study_design_scores_gemma":[0.000025228486,0.00085157884,0.9433787,0.000050198392,0.00030525902,0.00070285314,0.001658478,0.049065832,0.0016664536,0.001973433,0.0002798838,0.00004219166],"about_ca_topic_score_codex":0.0026380522,"about_ca_topic_score_gemma":0.0025579785,"teacher_disagreement_score":0.006945815,"about_ca_system_score_codex":0.0006648172,"about_ca_system_score_gemma":0.00067360606,"threshold_uncertainty_score":0.03673345},"labels":[],"label_agreement":null},{"id":"W4400406365","doi":"10.1007/s11334-024-00566-1","title":"Detecting mistakes in a domain model: a comparison of three approaches","year":2024,"lang":"en","type":"article","venue":"Innovations in Systems and Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science","score_opus":0.06713351668033267,"score_gpt":0.28202793810968657,"score_spread":0.2148944214293539,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400406365","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4443914,0.0032228674,0.52656025,0.0010262525,0.00023158928,0.00084132433,0.0014759387,0.011601571,0.010648837],"genre_scores_gemma":[0.7101226,0.0010609752,0.2823523,0.00029406205,0.000040071784,0.0001405832,0.0024990162,0.000574208,0.0029162415],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98745465,0.0035596397,0.0010052833,0.002195885,0.005008563,0.0007759054],"domain_scores_gemma":[0.95017356,0.030191554,0.0031844978,0.0059102043,0.0087575065,0.0017826934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074483687,0.0014625709,0.0015413167,0.006488992,0.0010934789,0.0039131856,0.0039222683,0.0026623327,0.002090036],"category_scores_gemma":[0.033962324,0.00057304424,0.001493182,0.0035767967,0.0009067095,0.0051334314,0.0041165138,0.0015956607,0.001144772],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004026343,0.0021615725,0.06412498,0.0016336942,0.001251703,0.0003524164,0.002986884,0.025422022,0.017109597,0.004810701,0.00710224,0.8690179],"study_design_scores_gemma":[0.00056668627,0.0020014385,0.07050119,0.00038041035,0.0015825059,0.0014394842,0.007488814,0.85001665,0.032079037,0.018337918,0.015281391,0.00032447462],"about_ca_topic_score_codex":0.012033155,"about_ca_topic_score_gemma":0.017038016,"teacher_disagreement_score":0.012033155,"about_ca_system_score_codex":0.0017506585,"about_ca_system_score_gemma":0.0032473768,"threshold_uncertainty_score":0.03939122},"labels":[],"label_agreement":null},{"id":"W4400421435","doi":"10.1016/j.procs.2024.06.069","title":"Triage Software Update Impact via Release Notes Classification","year":2024,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Triage; Categorization; Classifier (UML); Software; Machine learning; Process (computing); Artificial intelligence","score_opus":0.023406008543741343,"score_gpt":0.30285513759217464,"score_spread":0.2794491290484333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400421435","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8110056,0.0031748982,0.15041828,0.0019801138,0.0013187446,0.0016285286,0.011033577,0.01108447,0.008355836],"genre_scores_gemma":[0.90410054,0.0008762661,0.07979779,0.00026503537,0.00042417392,0.00027741268,0.009877244,0.00017877966,0.0042028017],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9973828,0.00042004013,0.00039954664,0.0005931715,0.0010190262,0.0001853516],"domain_scores_gemma":[0.9783392,0.009673639,0.005405713,0.0013931502,0.0043496247,0.0008387136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026257345,0.001368551,0.0007434028,0.0058304337,0.0005111127,0.0016601075,0.0013273628,0.00094184035,0.0020449103],"category_scores_gemma":[0.024385586,0.000250078,0.0005408855,0.0021965203,0.0002359254,0.0016850446,0.0010316004,0.0019360957,0.002371884],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001342749,0.0008456918,0.40001556,0.00048380098,0.00020639242,0.00074174715,0.00053454435,0.017389296,0.008553639,0.00059174874,0.02629297,0.5430019],"study_design_scores_gemma":[0.00013144047,0.0017608071,0.2725135,0.0003623108,0.0003044429,0.0019316546,0.0017590765,0.6748571,0.024329524,0.003086428,0.018727899,0.00023575677],"about_ca_topic_score_codex":0.0047125416,"about_ca_topic_score_gemma":0.0071696616,"teacher_disagreement_score":0.0058304337,"about_ca_system_score_codex":0.00064449286,"about_ca_system_score_gemma":0.0009105045,"threshold_uncertainty_score":0.013886392},"labels":[],"label_agreement":null},{"id":"W4400434361","doi":"10.1145/3715773","title":"An Adaptive Language-Agnostic Pruning Method for Greener Language Models for Code","year":2025,"lang":"en","type":"preprint","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Dalhousie University","funders":"Agencia Estatal de Investigación; Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada; Knut och Alice Wallenbergs Stiftelse","keywords":"Pruning; Computer science; Code (set theory); Artificial intelligence; Language model; Natural language processing; Programming language; Biology; Botany","score_opus":0.02915474167711867,"score_gpt":0.31616635034274054,"score_spread":0.28701160866562186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400434361","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02356315,0.00018692722,0.9668588,0.00013920067,0.000043022257,0.00007124526,0.00019660775,0.008108311,0.0008327282],"genre_scores_gemma":[0.24783222,0.0002408371,0.74167037,0.00035087785,0.000057169244,0.0002656317,0.0018483954,0.002052319,0.00568213],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99931073,0.00013638429,0.000040819566,0.00015996021,0.0002919869,0.00006009433],"domain_scores_gemma":[0.99846387,0.0007280737,0.00013902623,0.000291068,0.00033468675,0.00004329547],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00079750625,0.0008001346,0.00053734117,0.0010796604,0.00039378388,0.00077021134,0.0015387359,0.0006336298,0.0019760581],"category_scores_gemma":[0.0038108951,0.0003994316,0.0009833388,0.00063693535,0.00055192865,0.0014472075,0.0011911846,0.0013183923,0.0012535832],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003738739,0.00023456627,0.0038532645,0.00034539532,0.000117057476,0.00054117484,0.0005718721,0.14623164,0.09092842,0.016957905,0.012905958,0.72693884],"study_design_scores_gemma":[0.000021700365,0.00007223363,0.0004543065,0.00001708269,0.000024335008,0.00016673646,0.000053780062,0.96621424,0.021368982,0.006397401,0.0051913946,0.000017863144],"about_ca_topic_score_codex":0.0033419046,"about_ca_topic_score_gemma":0.009192916,"teacher_disagreement_score":0.0033419046,"about_ca_system_score_codex":0.00052977656,"about_ca_system_score_gemma":0.0013875167,"threshold_uncertainty_score":0.0066449046},"labels":[],"label_agreement":null},{"id":"W4400484631","doi":"10.1145/3663533.3664042","title":"Graph Neural Network vs. Large Language Model: A Comparative Analysis for Bug Report Priority and Severity Prediction","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Graph; Artificial neural network; Artificial intelligence; Natural language processing; Theoretical computer science","score_opus":0.02254936607785773,"score_gpt":0.3139234305496094,"score_spread":0.2913740644717517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400484631","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76465696,0.031622373,0.17193775,0.0046876213,0.00090152264,0.000298717,0.006520347,0.009902881,0.009471838],"genre_scores_gemma":[0.9551934,0.0036460923,0.031928655,0.00029494587,0.00014522445,0.00009578322,0.006237645,0.00022062044,0.0022377023],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988158,0.00046646493,0.00008744053,0.00026525446,0.0002616062,0.00010354569],"domain_scores_gemma":[0.9943408,0.0041195545,0.00029217044,0.00033048354,0.00076521473,0.00015171111],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003256663,0.0019119193,0.0010454466,0.003411271,0.00037629463,0.0009733234,0.0011989126,0.0010814592,0.00113816],"category_scores_gemma":[0.009404653,0.0003154345,0.0009942874,0.0021607298,0.0003715393,0.0022881178,0.0005696613,0.0013963425,0.0004944067],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014811393,0.00068962784,0.036926307,0.00080603827,0.0007337432,0.00026578218,0.00017803417,0.6014425,0.0017971746,0.0029297206,0.014726316,0.33802357],"study_design_scores_gemma":[0.000016608805,0.000118641816,0.0027252524,0.000024386021,0.00006815444,0.000025718482,0.00003964861,0.9946714,0.0003829458,0.0013157293,0.000596545,0.0000149823],"about_ca_topic_score_codex":0.03694808,"about_ca_topic_score_gemma":0.032491297,"teacher_disagreement_score":0.03694808,"about_ca_system_score_codex":0.0016782352,"about_ca_system_score_gemma":0.0013870562,"threshold_uncertainty_score":0.07346606},"labels":[],"label_agreement":null},{"id":"W4400484795","doi":"10.1145/3663529.3663835","title":"Unveil the Mystery of Critical Software Vulnerabilities","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Software; Secure coding; Software security assurance; Computer security; Programming language; Information security","score_opus":0.020895164527094657,"score_gpt":0.2953767885453353,"score_spread":0.2744816240182406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400484795","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8643747,0.020876821,0.03890874,0.01600683,0.00042378253,0.00015856323,0.04632943,0.004415745,0.008505393],"genre_scores_gemma":[0.8841758,0.006599108,0.033632748,0.0025941038,0.00055156177,0.00014075666,0.06918986,0.0010010551,0.0021151237],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9927166,0.0018592755,0.0008245585,0.001956664,0.002209107,0.00043375415],"domain_scores_gemma":[0.9613514,0.016138174,0.0098127695,0.007887427,0.0037163098,0.0010939109],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050970754,0.00080850336,0.0007032315,0.011044455,0.0011060085,0.003817621,0.0014781288,0.0015506542,0.0009093843],"category_scores_gemma":[0.042738408,0.00045672298,0.0007620437,0.010244235,0.0018863766,0.011313893,0.004536063,0.0019479452,0.00090936833],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036848345,0.00012831301,0.67958915,0.0016734838,0.0003804745,0.0011747482,0.005432234,0.0033916174,0.003590976,0.011760204,0.07381442,0.21869577],"study_design_scores_gemma":[0.00006824918,0.00020648581,0.6468209,0.0017839149,0.00029095015,0.0060634045,0.009499135,0.033157777,0.0057205902,0.0720562,0.22410202,0.00023033837],"about_ca_topic_score_codex":0.0053111115,"about_ca_topic_score_gemma":0.010371714,"teacher_disagreement_score":0.011044455,"about_ca_system_score_codex":0.0009509696,"about_ca_system_score_gemma":0.001275506,"threshold_uncertainty_score":0.02695626},"labels":[],"label_agreement":null},{"id":"W4400484804","doi":"10.1145/3663529.3663836","title":"Multi-line AI-Assisted Code Authoring","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Programming language; Line (geometry); Code (set theory)","score_opus":0.0635562543977823,"score_gpt":0.35022835168655614,"score_spread":0.28667209728877385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400484804","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22856547,0.00017278527,0.52555716,0.00055507425,0.00026086066,0.0004949954,0.00075253646,0.23152046,0.012120645],"genre_scores_gemma":[0.43910697,0.00010497652,0.5283468,0.00028946466,0.00007313413,0.0004920714,0.0014748324,0.014052119,0.016059648],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9970182,0.00091119396,0.00021748277,0.0006409581,0.0010230517,0.00018902548],"domain_scores_gemma":[0.9533701,0.025212007,0.0022167973,0.012412565,0.0050300695,0.0017584718],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035010173,0.0012891103,0.00064361753,0.001148646,0.00045041047,0.0016546289,0.0025732412,0.0012257302,0.011498953],"category_scores_gemma":[0.027858276,0.000879553,0.00063656026,0.000579541,0.0005111556,0.0032129318,0.0025098906,0.0020175234,0.005203727],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037965681,0.0018555637,0.015353229,0.0010171559,0.00018953536,0.001836734,0.007899096,0.018395653,0.18904345,0.0070885546,0.064934045,0.6885905],"study_design_scores_gemma":[0.00084712857,0.0024284162,0.015152833,0.0002509289,0.00016002484,0.0022878037,0.0012694604,0.63083255,0.14532495,0.010328503,0.19061734,0.00050009147],"about_ca_topic_score_codex":0.0007861468,"about_ca_topic_score_gemma":0.0017742849,"teacher_disagreement_score":0.011498953,"about_ca_system_score_codex":0.00047833714,"about_ca_system_score_gemma":0.00076267606,"threshold_uncertainty_score":0.038467824},"labels":[],"label_agreement":null},{"id":"W4400581753","doi":"10.1145/3660823","title":"Dependency-Induced Waste in Continuous Integration: An Empirical Study of Unused Dependencies in the npm Ecosystem","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Dependency (UML); Computer science; Context (archaeology); Reuse; JSON; Dependency theory (database theory); Code reuse; Resource (disambiguation); Dependency graph; Database; Software; Software engineering; Functional dependency; Relational database; Operating system; Engineering","score_opus":0.029745562251328024,"score_gpt":0.28998160021238495,"score_spread":0.26023603796105693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400581753","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9954235,0.00043760048,0.0013746375,0.00022463594,0.000010665648,0.000035494602,0.0010858712,0.00009767838,0.0013099182],"genre_scores_gemma":[0.9881016,0.00043923347,0.0040228893,0.00018023749,0.000026732361,0.00010651256,0.006213979,0.00015512391,0.0007537696],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9890295,0.0032735667,0.0011278585,0.0018212011,0.0038077154,0.0009401053],"domain_scores_gemma":[0.84726363,0.088366844,0.03421008,0.012808301,0.012979541,0.004371649],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011716055,0.0006501609,0.0006282326,0.005044703,0.0014175612,0.0028533752,0.002490155,0.0012516541,0.0015672402],"category_scores_gemma":[0.088709,0.0006617882,0.0007560916,0.0086644655,0.0017857152,0.006918644,0.003775374,0.0025943716,0.00082960894],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019988825,0.00050731545,0.9575647,0.00027509604,0.00012157528,0.00054835127,0.004235549,0.001932406,0.0004901989,0.001034956,0.0034376446,0.02965235],"study_design_scores_gemma":[0.000030139714,0.00034457105,0.9494114,0.00028441552,0.000120257755,0.0012076317,0.012000858,0.02017725,0.0010056028,0.0026381628,0.012696265,0.000083334846],"about_ca_topic_score_codex":0.008248563,"about_ca_topic_score_gemma":0.008294825,"teacher_disagreement_score":0.011716055,"about_ca_system_score_codex":0.0014982012,"about_ca_system_score_gemma":0.0018353411,"threshold_uncertainty_score":0.061961174},"labels":[],"label_agreement":null},{"id":"W4400581925","doi":"10.1145/3660813","title":"Revealing Software Development Work Patterns with PR-Issue Graph Topologies","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Network topology; Computer science; Graph; Software development; Software; Software engineering; Theoretical computer science; Programming language; Operating system","score_opus":0.014321233758022994,"score_gpt":0.23332791828836394,"score_spread":0.21900668453034094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400581925","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54024076,0.00089719024,0.4341388,0.0020073543,0.00006654034,0.00068969524,0.0036148457,0.002750587,0.015594171],"genre_scores_gemma":[0.74260086,0.00057677017,0.24983871,0.00010507762,0.000021845686,0.00046664482,0.0033443861,0.0004139714,0.0026317853],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99570256,0.002100476,0.00034330884,0.00065348257,0.0009875011,0.00021261531],"domain_scores_gemma":[0.95319545,0.034083117,0.004771584,0.004176746,0.0029716277,0.0008015468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039060013,0.00049113913,0.00032603703,0.009624139,0.0014128542,0.0034472833,0.0010890975,0.0009801278,0.002034848],"category_scores_gemma":[0.029285898,0.0005487078,0.00061628834,0.008820276,0.0012279552,0.0073525296,0.0028362991,0.0010475747,0.0005609393],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047901,0.00045526173,0.2787185,0.0024390935,0.00029052387,0.0039727157,0.15341642,0.03542673,0.018053902,0.12364497,0.014754814,0.36834803],"study_design_scores_gemma":[0.0001144702,0.00034538718,0.14020117,0.001169973,0.00032335953,0.004348526,0.12295408,0.25894177,0.020589061,0.27950794,0.1711908,0.00031345268],"about_ca_topic_score_codex":0.003086518,"about_ca_topic_score_gemma":0.008144623,"teacher_disagreement_score":0.009624139,"about_ca_system_score_codex":0.0011728029,"about_ca_system_score_gemma":0.0013577886,"threshold_uncertainty_score":0.020657122},"labels":[],"label_agreement":null},{"id":"W4400582230","doi":"10.1145/3660810","title":"ClarifyGPT: A Framework for Enhancing LLM-Based Code Generation via Requirements Clarification","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Natural Science Foundation of China","keywords":"Consistency (knowledge bases); Computer science; Fidelity; Code (set theory); Natural language generation; Natural language; Artificial intelligence; Programming language","score_opus":0.03864211965604995,"score_gpt":0.2956780983054859,"score_spread":0.25703597864943595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400582230","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033059258,0.00013131529,0.9238024,0.0003486057,0.000058803038,0.00056601514,0.00043907098,0.06945139,0.0018964295],"genre_scores_gemma":[0.036929715,0.00013779594,0.95135415,0.00037520428,0.000026763377,0.0006471363,0.0019025154,0.006736129,0.0018905431],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9939719,0.0024908965,0.0005322197,0.0010028702,0.0016720585,0.00033001867],"domain_scores_gemma":[0.9860563,0.0077466588,0.0012070166,0.0031370781,0.0013988545,0.00045412185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006709022,0.0025953315,0.00080138835,0.0023397529,0.0008374864,0.0020880944,0.0043767933,0.0024295985,0.008730297],"category_scores_gemma":[0.030963881,0.0017389818,0.0024571528,0.0009517119,0.0021650079,0.0038420695,0.004982915,0.004555638,0.004555242],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008521608,0.0010017274,0.0049293824,0.0030065512,0.0002532481,0.0016401716,0.004990538,0.07796579,0.06808623,0.07060592,0.08247399,0.6841943],"study_design_scores_gemma":[0.00043106094,0.0005191974,0.001237891,0.00045035666,0.000120660676,0.0011115681,0.0005765189,0.71627766,0.061271332,0.05199287,0.16573587,0.00027497933],"about_ca_topic_score_codex":0.005099506,"about_ca_topic_score_gemma":0.007589884,"teacher_disagreement_score":0.008730297,"about_ca_system_score_codex":0.0014769667,"about_ca_system_score_gemma":0.0039221975,"threshold_uncertainty_score":0.035481155},"labels":[],"label_agreement":null},{"id":"W4400582353","doi":"10.1145/3660809","title":"Mining Action Rules for Defect Reduction Planning","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Commit; Computer science; Counterfactual thinking; Reduction (mathematics); Precision and recall; Code (set theory); Action (physics); Recall; Compiler; Software; Baseline (sea); Machine learning; Software bug; Artificial intelligence; Software engineering; Programming language; Database; Set (abstract data type)","score_opus":0.03604877384191616,"score_gpt":0.29302717810681783,"score_spread":0.25697840426490165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400582353","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.116558895,0.0010209694,0.8617915,0.0014327086,0.00011318466,0.0006972683,0.0035132961,0.01190358,0.0029685507],"genre_scores_gemma":[0.51082784,0.00036824172,0.47935903,0.00034230965,0.000038686194,0.0005673634,0.006789538,0.00038841105,0.001318645],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99678504,0.00082464906,0.00029560697,0.0008252639,0.0010762104,0.0001932515],"domain_scores_gemma":[0.9869556,0.009423043,0.0010989328,0.0009850197,0.0013317445,0.00020571308],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024896304,0.0017791917,0.00095891056,0.0038234573,0.0007189459,0.0013245814,0.0018124526,0.0013567694,0.00186069],"category_scores_gemma":[0.0148372585,0.000628332,0.0020836003,0.0014467494,0.00089189596,0.0015581672,0.0011349677,0.0016161189,0.00071035617],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004092722,0.00073523924,0.04425475,0.0011454223,0.0003545457,0.0018752434,0.0008452783,0.37752467,0.013810488,0.009574661,0.010250832,0.5392195],"study_design_scores_gemma":[0.000050226983,0.00012675238,0.0022259147,0.000089882946,0.00011422155,0.00025834487,0.00019572588,0.97296745,0.007828868,0.012385721,0.003722615,0.00003426463],"about_ca_topic_score_codex":0.010321472,"about_ca_topic_score_gemma":0.017946163,"teacher_disagreement_score":0.010321472,"about_ca_system_score_codex":0.0012241732,"about_ca_system_score_gemma":0.003765405,"threshold_uncertainty_score":0.020522773},"labels":[],"label_agreement":null},{"id":"W4400582376","doi":"10.1145/3660807","title":"Do Large Language Models Pay Similar Attention Like Human Programmers When Generating Code?","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Programmer; Interpretability; Computer science; Code (set theory); Programming language; Code generation; Artificial intelligence; Computer security","score_opus":0.01965023609466952,"score_gpt":0.26783801853326755,"score_spread":0.24818778243859804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400582376","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70167685,0.0010546176,0.26437843,0.0054175216,0.00021407778,0.00017451786,0.0002560351,0.0065613403,0.020266566],"genre_scores_gemma":[0.96657914,0.00021170574,0.028446086,0.0015172684,0.00005028462,0.000067607485,0.00022581591,0.00097282726,0.0019292184],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99364907,0.0028781225,0.00019017086,0.0014672027,0.001347884,0.00046767507],"domain_scores_gemma":[0.9576057,0.027626192,0.0038080856,0.0071689407,0.0025850409,0.0012059471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005873315,0.00068798196,0.0006284129,0.0010121072,0.000537013,0.002671429,0.0012312421,0.0015363065,0.0027466267],"category_scores_gemma":[0.07296127,0.00070400804,0.0005386031,0.0006570037,0.0017723308,0.0056012277,0.001913059,0.0017322255,0.0012328412],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015269577,0.00050549547,0.22723801,0.001074846,0.0004929402,0.0014335606,0.041294463,0.01489269,0.11579852,0.02094825,0.015584166,0.5592101],"study_design_scores_gemma":[0.00046317242,0.0016478384,0.2861018,0.00048916374,0.0006633696,0.004270142,0.023148201,0.35037914,0.072976165,0.14630745,0.1130025,0.0005509808],"about_ca_topic_score_codex":0.0029032642,"about_ca_topic_score_gemma":0.0038196777,"teacher_disagreement_score":0.005873315,"about_ca_system_score_codex":0.00075460045,"about_ca_system_score_gemma":0.00091580744,"threshold_uncertainty_score":0.03106141},"labels":[],"label_agreement":null},{"id":"W4400582478","doi":"10.1145/3660793","title":"Towards Better Graph Neural Network-Based Fault Localization through Enhanced Code Representation","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; University of Alberta; Concordia University","funders":"","keywords":"Computer science; Debugging; Graph; Scalability; Inference; Software; Theoretical computer science; Leverage (statistics); Autoencoder; Artificial neural network; Software quality; Artificial intelligence; Data mining; Programming language; Software development","score_opus":0.019653216969188467,"score_gpt":0.27197463634407953,"score_spread":0.25232141937489105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400582478","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030130295,0.0001958221,0.96422476,0.00028862545,0.000034285826,0.000045168512,0.00031056494,0.00358698,0.0011834014],"genre_scores_gemma":[0.54199314,0.00038226056,0.4504933,0.00030437042,0.00004348388,0.0002128326,0.0023403696,0.0004211197,0.0038091645],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996983,0.00006373087,0.000016357471,0.00010004172,0.00008833451,0.000033362652],"domain_scores_gemma":[0.9991837,0.0003149828,0.00011629817,0.00014577463,0.00020740379,0.000031888772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00035818198,0.0009941454,0.00057570287,0.0015740021,0.0002441497,0.000738513,0.0014269209,0.00089097355,0.0018816788],"category_scores_gemma":[0.0026981542,0.0003532387,0.00071451237,0.0011138248,0.0005011826,0.0018087487,0.00079546607,0.0011954363,0.0005065162],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007591449,0.00007724834,0.0010232995,0.000089070745,0.000038793114,0.00007632711,0.000053414136,0.8097898,0.0062283427,0.0076752924,0.002531371,0.17234114],"study_design_scores_gemma":[0.0000024091394,0.0000072455578,0.00007042162,0.0000023681075,0.0000037078119,0.0000061636624,0.000003391206,0.9966917,0.00070041657,0.0022669842,0.00024329856,0.0000018665149],"about_ca_topic_score_codex":0.010944769,"about_ca_topic_score_gemma":0.01219986,"teacher_disagreement_score":0.010944769,"about_ca_system_score_codex":0.0010889341,"about_ca_system_score_gemma":0.00090454967,"threshold_uncertainty_score":0.021762133},"labels":[],"label_agreement":null},{"id":"W4400582781","doi":"10.1145/3643775","title":"Improving the Learning of Code Review Successive Tasks with Cross-Task Knowledge Distillation","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Task (project management); Code (set theory); Distillation; Artificial intelligence; Machine learning; Programming language; Engineering; Chemistry; Chromatography","score_opus":0.013152326303382412,"score_gpt":0.2774240932908377,"score_spread":0.2642717669874553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400582781","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2945325,0.0061196527,0.65123683,0.0021974158,0.001041046,0.00086999574,0.0012806783,0.0360985,0.0066233682],"genre_scores_gemma":[0.76029,0.0007936453,0.21960555,0.0015196475,0.00034219163,0.00051749457,0.0051514828,0.00091879175,0.010861157],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955831,0.0015128047,0.00027555777,0.001565814,0.00075423275,0.0003084509],"domain_scores_gemma":[0.9782507,0.011806264,0.0015301799,0.00292512,0.004623816,0.0008640044],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051596747,0.0024705927,0.0016729549,0.0019071148,0.00081820885,0.0017096024,0.0033676927,0.002824067,0.0023657097],"category_scores_gemma":[0.025589207,0.00072663365,0.001224848,0.0010752031,0.0009384231,0.0042698537,0.0031315458,0.0041490225,0.0018456471],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008641625,0.0010255239,0.0073400103,0.0008170148,0.0002262129,0.00026867623,0.0006103847,0.106452174,0.023082325,0.0014729665,0.021620931,0.8362196],"study_design_scores_gemma":[0.00009363012,0.00040024528,0.001384846,0.000045948058,0.0000744547,0.00011028905,0.00009769373,0.9759151,0.015545866,0.0028938218,0.0033904538,0.000047568825],"about_ca_topic_score_codex":0.0075472193,"about_ca_topic_score_gemma":0.012727662,"teacher_disagreement_score":0.0075472193,"about_ca_system_score_codex":0.0019552417,"about_ca_system_score_gemma":0.003066886,"threshold_uncertainty_score":0.027287304},"labels":[],"label_agreement":null},{"id":"W4400582900","doi":"10.1145/3643774","title":"AI-Assisted Code Authoring at Scale: Fine-Tuning, Deploying, and Mixed Methods Evaluation","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Code (set theory); Scale (ratio); Programming language; Geography; Cartography","score_opus":0.033052256068704086,"score_gpt":0.3336914157279161,"score_spread":0.300639159659212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400582900","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5437647,0.0032712661,0.38323352,0.0024563,0.0007062853,0.0031225628,0.0023453876,0.045949865,0.01515006],"genre_scores_gemma":[0.53281695,0.00039613416,0.4541804,0.00075457914,0.000083789004,0.002540886,0.003033073,0.0038824757,0.0023117452],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98529196,0.009986321,0.0009595599,0.0014715168,0.0018485897,0.0004419814],"domain_scores_gemma":[0.86084145,0.11359395,0.0020227758,0.013958023,0.0077605206,0.0018232971],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02933411,0.0019187815,0.0007821088,0.001645065,0.0009150905,0.0021645639,0.003863279,0.0021455523,0.0029491577],"category_scores_gemma":[0.103697516,0.00089986477,0.0012223992,0.0010854851,0.0019176753,0.0029183433,0.0039117327,0.0034613619,0.0013994803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0045553236,0.005040111,0.0458971,0.0039348523,0.0011117785,0.00064866175,0.007724445,0.24294372,0.022447033,0.016244328,0.034470096,0.6149826],"study_design_scores_gemma":[0.0012432722,0.0014389094,0.006654484,0.00037230502,0.00022238432,0.0001537621,0.0012140416,0.9453399,0.013294703,0.013540263,0.01639692,0.00012901807],"about_ca_topic_score_codex":0.0070797163,"about_ca_topic_score_gemma":0.008663188,"teacher_disagreement_score":0.02933411,"about_ca_system_score_codex":0.0016353606,"about_ca_system_score_gemma":0.002205478,"threshold_uncertainty_score":0.15513545},"labels":[],"label_agreement":null},{"id":"W4400647264","doi":"10.1109/tse.2024.3428324","title":"Towards Efficient Fine-Tuning of Language Models With Organizational Data for Automated Software Review","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Software engineering; Software; Programming language; Data science","score_opus":0.024687243253182226,"score_gpt":0.27963316747094596,"score_spread":0.25494592421776374,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400647264","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09542399,0.0037277166,0.86859936,0.001632009,0.000241258,0.00057171885,0.0010892415,0.026985167,0.0017296075],"genre_scores_gemma":[0.5518605,0.0007948928,0.4355762,0.0013502813,0.00019006156,0.0007630682,0.0052900068,0.00094037026,0.003234671],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945503,0.002802685,0.0003320625,0.0013313191,0.00074815744,0.0002355507],"domain_scores_gemma":[0.98044306,0.012307264,0.0015245636,0.0020939456,0.0028496187,0.00078160304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059137875,0.0016712563,0.0016724875,0.0026833836,0.0006740881,0.0018034085,0.0035014765,0.0021220036,0.0015303778],"category_scores_gemma":[0.032909315,0.00089931017,0.0014978289,0.001310255,0.0008327045,0.004214533,0.002694579,0.0036359848,0.001743206],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000580661,0.0009776854,0.012274203,0.0010960257,0.00033175544,0.00033283286,0.0010044737,0.21396476,0.023921054,0.0030923332,0.02247512,0.719949],"study_design_scores_gemma":[0.000044353383,0.000118338416,0.000690258,0.000033228305,0.000030191855,0.000062409345,0.00008285354,0.98906404,0.0046251835,0.0030903295,0.0021356046,0.000023206057],"about_ca_topic_score_codex":0.009124827,"about_ca_topic_score_gemma":0.018876992,"teacher_disagreement_score":0.009124827,"about_ca_system_score_codex":0.0017189933,"about_ca_system_score_gemma":0.004149133,"threshold_uncertainty_score":0.03127551},"labels":[],"label_agreement":null},{"id":"W4400647844","doi":"10.1016/j.cose.2024.103994","title":"SCL-CVD: Supervised contrastive learning for code vulnerability detection via GraphCodeBERT","year":2024,"lang":"en","type":"article","venue":"Computers & Security","topic":"Software Engineering Research","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Vulnerability (computing); Code (set theory); Artificial intelligence; Computer security; Programming language","score_opus":0.014764084794374605,"score_gpt":0.2703110922352167,"score_spread":0.25554700744084213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400647844","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051431954,0.00071237644,0.9074671,0.00036837716,0.00018778046,0.00031834663,0.0023876135,0.033878427,0.0032480028],"genre_scores_gemma":[0.34715113,0.00026996803,0.63052976,0.0005560736,0.00014940037,0.00044855924,0.009769035,0.0022246912,0.008901278],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989962,0.00017438906,0.00003576135,0.00042301184,0.00025793098,0.00011273953],"domain_scores_gemma":[0.9978357,0.0009896223,0.00015547172,0.0004943633,0.00040563932,0.00011927719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010498911,0.0019253257,0.0015599731,0.0036953343,0.00090886763,0.0010211699,0.0034685142,0.0021641664,0.0051905257],"category_scores_gemma":[0.0037263862,0.00069455995,0.0014098932,0.0016772945,0.0010226137,0.0018426016,0.002279737,0.0029719523,0.0028026216],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058615016,0.00070354925,0.0044709463,0.0004094771,0.00022591828,0.0002795938,0.00011784564,0.07634179,0.030222293,0.006744183,0.03996397,0.83993423],"study_design_scores_gemma":[0.000038359387,0.000103054954,0.00069298805,0.000018604236,0.0000253076,0.00008343266,0.000022562303,0.9803695,0.007832351,0.007555044,0.0032387166,0.000020132135],"about_ca_topic_score_codex":0.009070359,"about_ca_topic_score_gemma":0.022827843,"teacher_disagreement_score":0.009070359,"about_ca_system_score_codex":0.0010780056,"about_ca_system_score_gemma":0.0016516338,"threshold_uncertainty_score":0.018035114},"labels":[],"label_agreement":null},{"id":"W4400652308","doi":"10.1145/3678172","title":"A Disruptive Research Playbook for Studying Disruptive Innovations","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Disruptive technology; Transformative learning; Framing (construction); Disruptive innovation; Emerging technologies; Software; Generative grammar; Data science; Empirical research; Management science; Artificial intelligence; Sociology; Engineering","score_opus":0.25035296368953364,"score_gpt":0.4320521837769804,"score_spread":0.18169922008744677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400652308","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032010578,0.020807981,0.5353144,0.06540525,0.008279945,0.0055271643,0.0040495098,0.0029728964,0.32563233],"genre_scores_gemma":[0.223963,0.033452153,0.559148,0.011059996,0.0025861235,0.014850823,0.0026357123,0.0013393711,0.15096486],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99417514,0.004287645,0.00023527641,0.00041565823,0.00062762195,0.00025872953],"domain_scores_gemma":[0.9786702,0.01689458,0.00064651197,0.0016307962,0.001206871,0.0009508913],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.007076492,0.0029967115,0.001107724,0.0051207766,0.0058897953,0.008435646,0.0030388837,0.004372289,0.017917126],"category_scores_gemma":[0.014425104,0.0006947469,0.0010226439,0.004755674,0.012185089,0.012465757,0.0064028017,0.007711138,0.0040725255],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000111675785,0.0002717613,0.0008803118,0.0015082475,0.00001705177,0.00092363555,0.07051655,0.001693651,0.0029445544,0.70100474,0.09927658,0.12085127],"study_design_scores_gemma":[0.000040918403,0.00025790883,0.0008660795,0.0018784575,0.000013016168,0.0007463651,0.04019135,0.0019473154,0.0012457025,0.10310239,0.8496388,0.00007161091],"about_ca_topic_score_codex":0.0056974245,"about_ca_topic_score_gemma":0.010940088,"teacher_disagreement_score":0.9929235,"about_ca_system_score_codex":0.00559663,"about_ca_system_score_gemma":0.0063396064,"threshold_uncertainty_score":0.05993879},"labels":[],"label_agreement":null},{"id":"W4400680686","doi":"10.1109/saner60148.2024.00065","title":"TraceJIT: Evaluating the Impact of Behavioral Code Change on Just-In-Time Defect Prediction","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Japan Society for the Promotion of Science","keywords":"Computer science; Code (set theory); Programming language","score_opus":0.13524793939442578,"score_gpt":0.4260565789846571,"score_spread":0.2908086395902313,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400680686","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94856864,0.0018473773,0.033236843,0.0006335032,0.00019250039,0.00013544611,0.003918101,0.00908025,0.0023874165],"genre_scores_gemma":[0.96475154,0.00024552524,0.026360491,0.00011228236,0.00005925357,0.00006160128,0.007165059,0.00025305199,0.0009912252],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.99690336,0.0009273878,0.00022344706,0.0007666843,0.0009856541,0.00019352397],"domain_scores_gemma":[0.9684664,0.021197999,0.0026897518,0.0036882965,0.002735783,0.0012216818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060109273,0.0017452178,0.0007554485,0.0027968574,0.00036048246,0.0013694703,0.0014016832,0.0013648071,0.0009820056],"category_scores_gemma":[0.02937843,0.00035047886,0.0010051491,0.0012797931,0.0006591844,0.0023248342,0.0010935012,0.0017445044,0.00054025307],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015973428,0.0014333318,0.33963764,0.00046929868,0.00065645087,0.00025883096,0.00029970362,0.4502912,0.004844302,0.0012158725,0.009654321,0.18964167],"study_design_scores_gemma":[0.00004517737,0.0008683685,0.02518089,0.000034303837,0.00006801721,0.00010017591,0.000076591234,0.9691862,0.0022088715,0.00096603745,0.0012240377,0.00004135481],"about_ca_topic_score_codex":0.015569937,"about_ca_topic_score_gemma":0.013935352,"teacher_disagreement_score":0.015569937,"about_ca_system_score_codex":0.0010309293,"about_ca_system_score_gemma":0.0014686855,"threshold_uncertainty_score":0.031789243},"labels":[],"label_agreement":null},{"id":"W4400681025","doi":"10.1109/saner60148.2024.00083","title":"Can We Identify Stack Overflow Questions Requiring Code Snippets? Investigating the Cause &amp; Effect of Missing Code Snippets","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; University of Saskatchewan","funders":"","keywords":"Computer science; Code (set theory); Stack (abstract data type); Programming language; Information retrieval; Data mining","score_opus":0.04726092849989341,"score_gpt":0.34931407322824537,"score_spread":0.302053144728352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400681025","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98018265,0.0008795919,0.011333372,0.0014276601,0.0000699785,0.00044685233,0.0013614652,0.0011483086,0.0031500645],"genre_scores_gemma":[0.98505026,0.0003002012,0.0114858495,0.0003813621,0.000043478925,0.000274532,0.0010969839,0.00019399053,0.0011734106],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9762822,0.0092968885,0.002819081,0.0031784864,0.0070134946,0.0014098122],"domain_scores_gemma":[0.4904086,0.37110972,0.090806864,0.011001223,0.03215663,0.0045168963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020127067,0.0010135175,0.00084822386,0.005538208,0.0009990956,0.0033842192,0.001128174,0.0023087754,0.0036252146],"category_scores_gemma":[0.28056678,0.0008073797,0.0010968324,0.0030614415,0.0015570646,0.0062943804,0.0025139782,0.0019708907,0.0014119396],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010536767,0.0005820559,0.8837473,0.002112581,0.00020278853,0.0013086682,0.021928264,0.0009427818,0.0067038485,0.0009593291,0.0031857751,0.07727287],"study_design_scores_gemma":[0.000067038105,0.0012286312,0.93828964,0.0008355341,0.0003166872,0.001490872,0.01867017,0.014912966,0.011039884,0.002366851,0.010587414,0.00019428742],"about_ca_topic_score_codex":0.0046950444,"about_ca_topic_score_gemma":0.006105673,"teacher_disagreement_score":0.020127067,"about_ca_system_score_codex":0.0014681477,"about_ca_system_score_gemma":0.0031116032,"threshold_uncertainty_score":0.106443405},"labels":[],"label_agreement":null},{"id":"W4400681240","doi":"10.1109/saner60148.2024.00009","title":"On the Prevalence, Co-occurrence, and Impact of Infrastructure-as-Code Smells","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Computer science; Code (set theory); Code smell; Computer security; Programming language; Software; Software quality","score_opus":0.013257333416513072,"score_gpt":0.312721260720004,"score_spread":0.2994639273034909,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400681240","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99707186,0.00036687194,0.0011382177,0.00015361104,0.0000067269452,0.00002711869,0.00011692176,0.000046596964,0.0010720455],"genre_scores_gemma":[0.99837095,0.00019825532,0.00093314284,0.000033328968,0.000009323849,0.000022900964,0.00018143065,0.000024825073,0.0002258992],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9761431,0.0063860063,0.0024212948,0.0035371734,0.009919888,0.0015925007],"domain_scores_gemma":[0.69262546,0.17084461,0.09581474,0.010334561,0.024464026,0.005916572],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014597997,0.0005550724,0.0005284164,0.0072715087,0.0009399304,0.0030904263,0.0010332587,0.0012671925,0.0015549703],"category_scores_gemma":[0.112531364,0.00052280555,0.0005832098,0.0046016783,0.0025198308,0.005137924,0.004135908,0.001450644,0.0003755696],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010514861,0.00009625212,0.9698467,0.00016746134,0.0000643736,0.00049200904,0.0075245295,0.00033137717,0.0010623402,0.00022149719,0.00023853367,0.01984975],"study_design_scores_gemma":[0.0000037158823,0.00014782633,0.9889748,0.000121253324,0.000034238074,0.00060774194,0.0074105156,0.001175788,0.0005130877,0.00023978548,0.0007389989,0.00003227265],"about_ca_topic_score_codex":0.0041297516,"about_ca_topic_score_gemma":0.007348076,"teacher_disagreement_score":0.014597997,"about_ca_system_score_codex":0.0012517691,"about_ca_system_score_gemma":0.001297269,"threshold_uncertainty_score":0.0772025},"labels":[],"label_agreement":null},{"id":"W4400681481","doi":"10.1109/saner60148.2024.00035","title":"Investigating and Detecting Silent Bugs in PyTorch Programs","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Software bug; Programming language; Artificial intelligence; Software","score_opus":0.029496780471090576,"score_gpt":0.2851109023873632,"score_spread":0.25561412191627264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400681481","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.626317,0.0010797529,0.23960079,0.00078809116,0.0001606921,0.0004373929,0.0057538515,0.12236966,0.0034927141],"genre_scores_gemma":[0.76805305,0.0004086728,0.21380231,0.00046467205,0.000052403637,0.00030232503,0.011033488,0.0023123731,0.0035707103],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9954283,0.0009160188,0.00042876572,0.0011829714,0.0016824377,0.0003615426],"domain_scores_gemma":[0.9794526,0.011917736,0.0033838933,0.0021906393,0.0026502362,0.0004048648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035490894,0.0014563571,0.00069945847,0.004434805,0.00057908846,0.0010818435,0.0017507534,0.0011382331,0.001261991],"category_scores_gemma":[0.0234517,0.0005843934,0.0007738773,0.0014904852,0.0007666558,0.0024900886,0.002022881,0.0013044173,0.0010228436],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088495377,0.0007427647,0.21371399,0.002131776,0.00028310608,0.0028747965,0.003047228,0.02649195,0.057265192,0.0026754658,0.042265013,0.6476237],"study_design_scores_gemma":[0.0001427123,0.00055502483,0.08127768,0.0003455139,0.00020793629,0.0017504907,0.0010243752,0.77379644,0.10896398,0.0062543456,0.025516866,0.00016458436],"about_ca_topic_score_codex":0.004131354,"about_ca_topic_score_gemma":0.008934091,"teacher_disagreement_score":0.004434805,"about_ca_system_score_codex":0.0006924031,"about_ca_system_score_gemma":0.001754876,"threshold_uncertainty_score":0.018769562},"labels":[],"label_agreement":null},{"id":"W4401047029","doi":"10.1016/j.jss.2024.112159","title":"EvaluateXAI: A framework to evaluate the reliability and consistency of rule-based XAI techniques for software analytics tasks","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Consistency (knowledge bases); Computer science; Analytics; Reliability (semiconductor); Data mining; Software; Reliability engineering; Artificial intelligence; Engineering; Programming language","score_opus":0.03221025691242322,"score_gpt":0.32606641202743425,"score_spread":0.293856155115011,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401047029","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09570283,0.0024342379,0.86313915,0.0003581855,0.0002285949,0.0010886077,0.0021803186,0.031801242,0.003066797],"genre_scores_gemma":[0.31408706,0.00039733393,0.68036216,0.00013029923,0.00008630934,0.0007898395,0.0022845666,0.0010205732,0.00084179715],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9748512,0.0074460367,0.0024087294,0.0023243437,0.012076462,0.0008932458],"domain_scores_gemma":[0.93447155,0.035091635,0.007561355,0.010965149,0.010914529,0.0009958097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028527722,0.0026074639,0.0018192433,0.0073891305,0.0010697219,0.004602469,0.005949168,0.0023283595,0.0019402639],"category_scores_gemma":[0.08272753,0.0010534568,0.00192064,0.0030714558,0.0014559951,0.005206534,0.0035906788,0.0029906002,0.000870175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036426473,0.0016609496,0.08367467,0.0022232977,0.0024526524,0.00033870104,0.0008550838,0.22181793,0.027893413,0.01828221,0.010908736,0.62624973],"study_design_scores_gemma":[0.00023853453,0.0012625447,0.011206587,0.00019257923,0.00026916963,0.00022161887,0.00018161502,0.9501271,0.019918615,0.012841335,0.0034124763,0.00012770135],"about_ca_topic_score_codex":0.005441378,"about_ca_topic_score_gemma":0.004867177,"teacher_disagreement_score":0.028527722,"about_ca_system_score_codex":0.0013803777,"about_ca_system_score_gemma":0.0034138625,"threshold_uncertainty_score":0.1508708},"labels":[],"label_agreement":null},{"id":"W4401078640","doi":"10.1145/3643691.3648586","title":"Code Ownership in Open-Source AI Software Security","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Open source software; Programming language; Open source; Source code; Code (set theory); Software engineering; Code review; Software; Computer security; Static program analysis; Software development","score_opus":0.028412806578632575,"score_gpt":0.3093841385384059,"score_spread":0.2809713319597733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401078640","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98917705,0.00016861645,0.008111512,0.0002031989,0.000005811442,0.00004232955,0.00008316265,0.00009115436,0.0021172015],"genre_scores_gemma":[0.9978757,0.000043199252,0.0017277362,0.000012363509,0.00000491187,0.000029071836,0.000057419686,0.000030788902,0.00021884979],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9867866,0.0045572766,0.0011356093,0.0014893473,0.00520369,0.0008274602],"domain_scores_gemma":[0.752564,0.119832516,0.085473195,0.017415676,0.01994306,0.0047716815],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01331408,0.00041603358,0.0002461662,0.005533577,0.0009235668,0.002642672,0.0007118901,0.0005987136,0.0015069599],"category_scores_gemma":[0.1284503,0.00040357347,0.0003505056,0.0038095668,0.0040610014,0.006505863,0.005481481,0.0010297697,0.00023881138],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016167272,0.00011839605,0.94213355,0.000098619734,0.00005579784,0.00021167539,0.005758605,0.003536061,0.002144167,0.004120813,0.0003099831,0.04135067],"study_design_scores_gemma":[0.000009851949,0.00023049572,0.9629831,0.00014210375,0.000042321713,0.0004904877,0.0036572886,0.021503838,0.0028279885,0.00610029,0.0019598675,0.000052426873],"about_ca_topic_score_codex":0.0035033887,"about_ca_topic_score_gemma":0.0036835575,"teacher_disagreement_score":0.01331408,"about_ca_system_score_codex":0.0017089385,"about_ca_system_score_gemma":0.0016924749,"threshold_uncertainty_score":0.07041246},"labels":[],"label_agreement":null},{"id":"W4401123839","doi":"10.1145/3680472","title":"MULTICR: Predicting Merged and Abandoned Code Changes in Modern Code Review Using Multi-Objective Search","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Code review; Code (set theory); Machine learning; Eclipse; Process (computing); Artificial intelligence; Software quality; Interoperability; Genetic programming; Search-based software engineering; Data mining; Software engineering; Software; Software development; Software development process; Programming language; World Wide Web","score_opus":0.15337003444371033,"score_gpt":0.38079161845485976,"score_spread":0.22742158401114942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401123839","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7379465,0.0044424627,0.24626245,0.0012534689,0.00016908845,0.00069640455,0.0016541074,0.0046809195,0.002894548],"genre_scores_gemma":[0.85702187,0.00048803622,0.1371533,0.00032063632,0.00007277452,0.00026496744,0.0024729182,0.00015444114,0.0020511039],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99744713,0.0008908871,0.00024423184,0.00069322495,0.00055141357,0.00017314203],"domain_scores_gemma":[0.9885583,0.0075792097,0.0015108705,0.00044560884,0.0015469706,0.00035900425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005171971,0.0015234766,0.0013893448,0.004711947,0.0004383468,0.0013590708,0.0018752705,0.0016202342,0.0011421992],"category_scores_gemma":[0.014349556,0.0004944665,0.0011368623,0.0024626448,0.0004267539,0.0013367286,0.0011587875,0.0011103225,0.00037647886],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064254616,0.0009202978,0.1053193,0.0012504554,0.00074710196,0.0006024982,0.0006521185,0.47775835,0.0045683547,0.0019979945,0.008971931,0.39656898],"study_design_scores_gemma":[0.000031263404,0.00020441962,0.0061316905,0.000046130084,0.000060869817,0.000070824484,0.00007928223,0.99080795,0.0010672525,0.0008202875,0.00065950595,0.000020596268],"about_ca_topic_score_codex":0.011535248,"about_ca_topic_score_gemma":0.018510027,"teacher_disagreement_score":0.011535248,"about_ca_system_score_codex":0.0012509265,"about_ca_system_score_gemma":0.0023283206,"threshold_uncertainty_score":0.027352333},"labels":[],"label_agreement":null},{"id":"W4401161820","doi":"10.1587/transinf.2023edp7238","title":"Unveiling Python Version Compatibility Challenges in Code Snippets on Stack Overflow","year":2024,"lang":"en","type":"article","venue":"IEICE Transactions on Information and Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Python (programming language); Computer science; Programming language; Source code; Compatibility (geochemistry); World Wide Web; Engineering","score_opus":0.04424765382061606,"score_gpt":0.2838154316304337,"score_spread":0.23956777780981764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401161820","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85370415,0.00090510544,0.109052785,0.00407942,0.0005393898,0.00073903945,0.003950374,0.009451562,0.017578175],"genre_scores_gemma":[0.8694469,0.000716204,0.110025205,0.002094866,0.00024863417,0.00052382814,0.0042823167,0.0047490806,0.007913017],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98574674,0.005146487,0.0013475621,0.0013449641,0.0058847643,0.00052944897],"domain_scores_gemma":[0.8205576,0.13300882,0.020160442,0.011032945,0.014047841,0.0011923503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012357725,0.0009761587,0.00057080865,0.0043686684,0.0019029694,0.0038807713,0.0014282444,0.0016169263,0.00416589],"category_scores_gemma":[0.13599744,0.00067184743,0.00063524756,0.003456313,0.002521953,0.009753396,0.005054406,0.0025697532,0.0011007393],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022958645,0.00087141345,0.25471434,0.00537946,0.00033271904,0.013716389,0.10509417,0.010381183,0.030535897,0.036209773,0.050463304,0.49000543],"study_design_scores_gemma":[0.00024067248,0.0016917243,0.27099708,0.009831312,0.0007199015,0.018136932,0.071010314,0.09498221,0.09838593,0.10088555,0.33187306,0.0012453097],"about_ca_topic_score_codex":0.002343081,"about_ca_topic_score_gemma":0.0057258336,"teacher_disagreement_score":0.012357725,"about_ca_system_score_codex":0.001157649,"about_ca_system_score_gemma":0.0024314066,"threshold_uncertainty_score":0.065354705},"labels":[],"label_agreement":null},{"id":"W4401177203","doi":"10.1007/s00521-024-10131-3","title":"Optimizing beyond boundaries: empowering the salp swarm algorithm for global optimization and defective software module classification","year":2024,"lang":"en","type":"article","venue":"Neural Computing and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Ajman University","keywords":"Computational Science and Engineering; Computer science; Swarm behaviour; Software; Algorithm; Optimization algorithm; Artificial intelligence; Mathematical optimization; Machine learning; Mathematics","score_opus":0.01722681091511139,"score_gpt":0.3012608134556873,"score_spread":0.28403400254057587,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401177203","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025031002,0.00009486807,0.972345,0.00016324791,0.000027269942,0.0000212908,0.000018563454,0.00052808406,0.0017706829],"genre_scores_gemma":[0.5976356,0.00008800732,0.39923865,0.00019049304,0.00003974059,0.00009421172,0.00013231496,0.000274392,0.0023065913],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995627,0.00011792087,0.000025602332,0.00008369725,0.0001621525,0.000047909838],"domain_scores_gemma":[0.99894696,0.00053913885,0.00009738251,0.00015319674,0.00022013862,0.00004306672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010391928,0.0007074977,0.0010204882,0.0006848756,0.0004119172,0.0010927307,0.0012056315,0.0012790699,0.0013551184],"category_scores_gemma":[0.0045317328,0.00036321318,0.00063912076,0.00047883755,0.00082401995,0.0012846532,0.0013718401,0.0010847693,0.00039173002],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006331471,0.000046938017,0.0010834174,0.000045924877,0.00002999739,0.000068683104,0.00012469517,0.86057365,0.0050774734,0.009766779,0.0016402592,0.1214788],"study_design_scores_gemma":[0.000001630681,0.0000072669714,0.000033464847,0.0000017458154,0.0000013935138,0.0000048037064,0.000004471918,0.9972825,0.0003376244,0.0021993213,0.00012449622,0.0000012276062],"about_ca_topic_score_codex":0.0030775166,"about_ca_topic_score_gemma":0.002946905,"teacher_disagreement_score":0.0030775166,"about_ca_system_score_codex":0.00046526766,"about_ca_system_score_gemma":0.0008371231,"threshold_uncertainty_score":0.0061191916},"labels":[],"label_agreement":null},{"id":"W4401253331","doi":"10.1007/s10664-024-10514-z","title":"IRJIT: A simple, online, information retrieval approach for just-in-time software defect prediction","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Alberta","funders":"","keywords":"Computer science; Simple (philosophy); Information retrieval; Data mining; Software; Machine learning; Artificial intelligence; Programming language","score_opus":0.02499642994780278,"score_gpt":0.28696472697798037,"score_spread":0.2619682970301776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401253331","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024229,0.0014116683,0.9048358,0.0002769655,0.0003829768,0.0005856165,0.0057246634,0.058323096,0.004230149],"genre_scores_gemma":[0.17708942,0.00081049284,0.800072,0.0003234922,0.00042125955,0.00055120914,0.010663247,0.0011435831,0.008925277],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977552,0.000421977,0.00020515954,0.0003883019,0.0010816557,0.00014773577],"domain_scores_gemma":[0.9958358,0.0017591261,0.00038779984,0.0009405782,0.0009013762,0.00017522903],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002178702,0.002038982,0.002214526,0.0060680304,0.0006905059,0.0020196203,0.0033612442,0.0017824145,0.006242049],"category_scores_gemma":[0.009583185,0.00051715784,0.0014347476,0.0037278612,0.0004088961,0.00387443,0.0019290586,0.001532655,0.0062350156],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007810759,0.0010545466,0.0058878404,0.0007455108,0.00036831884,0.0003364351,0.00012450937,0.014620537,0.02048379,0.0041927164,0.0443754,0.90702933],"study_design_scores_gemma":[0.00019887743,0.00079149334,0.005880931,0.0000747267,0.00029111482,0.00089470146,0.000121027406,0.9249217,0.031626076,0.014328602,0.020663936,0.0002067572],"about_ca_topic_score_codex":0.0039370563,"about_ca_topic_score_gemma":0.0064925184,"teacher_disagreement_score":0.006242049,"about_ca_system_score_codex":0.00050859334,"about_ca_system_score_gemma":0.0014879048,"threshold_uncertainty_score":0.020881712},"labels":[],"label_agreement":null},{"id":"W4401359429","doi":"10.1016/j.jss.2024.112179","title":"An empirical study on bug severity estimation using source code metrics and static analysis","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Manitoba; University of Calgary","funders":"","keywords":"Software bug; Computer science; Leverage (statistics); Java; Source lines of code; Code (set theory); Source code; Open source; Static program analysis; Static analysis; Set (abstract data type); Data mining; Code review; Domain (mathematical analysis); Software; Categorization; Machine learning; Artificial intelligence; Programming language; Software development; Mathematics","score_opus":0.04805113706718376,"score_gpt":0.36145082974943565,"score_spread":0.3133996926822519,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401359429","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99337125,0.0011063191,0.0025375811,0.00023632341,0.00001599278,0.000038924718,0.001644973,0.00013822824,0.00091042],"genre_scores_gemma":[0.9914954,0.00037629952,0.0029970498,0.00005935441,0.000033338147,0.00004621764,0.0046328492,0.000057575307,0.00030204497],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98892397,0.0044419784,0.0013261347,0.0017786892,0.00311793,0.00041133622],"domain_scores_gemma":[0.74972135,0.19625385,0.026007788,0.009412655,0.016422158,0.0021822366],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012299812,0.00080245786,0.00061493536,0.0056962683,0.0004973111,0.0012531207,0.0010263671,0.00094272825,0.0007512004],"category_scores_gemma":[0.09635544,0.00041096684,0.000616079,0.0047463328,0.00087365194,0.0032127714,0.0009425765,0.0013946838,0.00046488259],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018517913,0.00031167956,0.95818347,0.00031554003,0.00017644979,0.00015191828,0.00087483233,0.0024422957,0.00076415983,0.0002478798,0.0021203812,0.03422624],"study_design_scores_gemma":[0.000031604228,0.00045481775,0.94908005,0.00018501042,0.00010815393,0.00084817153,0.0014525735,0.04123133,0.0014434297,0.0005827763,0.0045330888,0.00004907724],"about_ca_topic_score_codex":0.0037190716,"about_ca_topic_score_gemma":0.004214601,"teacher_disagreement_score":0.012299812,"about_ca_system_score_codex":0.0005687231,"about_ca_system_score_gemma":0.0005350668,"threshold_uncertainty_score":0.0650484},"labels":[],"label_agreement":null},{"id":"W4401414416","doi":"10.1109/ms.2024.3440190","title":"A State-of-the-Practice Release-Readiness Checklist for Generative AI-Based Software Products: A Gray Literature Survey","year":2024,"lang":"en","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bank of Canada; National Bank of Canada; Business Development Bank of Canada; Queen's University","funders":"","keywords":"Software engineering; Checklist; Computer science; Software peer review; Software development; Software; Generative grammar; State (computer science); Software release life cycle; Software quality; Engineering management; Software construction; Engineering; Programming language; Artificial intelligence; Psychology","score_opus":0.01837307102445155,"score_gpt":0.2940337656147816,"score_spread":0.27566069459033005,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401414416","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17985384,0.3839487,0.28359658,0.051315926,0.001417899,0.00435869,0.0030599325,0.00632307,0.08612535],"genre_scores_gemma":[0.45552966,0.21007925,0.31578788,0.004388963,0.00038084065,0.003384261,0.004760479,0.0010341982,0.004654411],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95813644,0.012798192,0.010152203,0.0023634988,0.015229698,0.0013199772],"domain_scores_gemma":[0.6911216,0.18449232,0.028760472,0.0118419705,0.07917626,0.004607349],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.051211778,0.001179579,0.0015069864,0.018634384,0.0022250717,0.007846904,0.004282611,0.002523297,0.00253115],"category_scores_gemma":[0.17592722,0.0011908635,0.001196628,0.010371754,0.003730793,0.011180761,0.004054109,0.002715519,0.001987969],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013295958,0.00037126848,0.021381894,0.016484756,0.00010250167,0.00023534017,0.00677727,0.0021150922,0.0023376842,0.016337723,0.017198978,0.9165245],"study_design_scores_gemma":[0.000120137745,0.0027554843,0.09847998,0.16179945,0.00094165956,0.003162283,0.047550775,0.020456487,0.019017043,0.04237653,0.60240734,0.0009329029],"about_ca_topic_score_codex":0.0070310854,"about_ca_topic_score_gemma":0.010481186,"teacher_disagreement_score":0.9487882,"about_ca_system_score_codex":0.00412288,"about_ca_system_score_gemma":0.018948903,"threshold_uncertainty_score":0.270837},"labels":[],"label_agreement":null},{"id":"W4401454847","doi":"10.1145/3643664.3648202","title":"Emerging Results on Automated Support for Searching and Selecting Evidence for Systematic Literature Review Updates","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Chicoutimi","funders":"","keywords":"Computer science; Data science; Systematic review; Information retrieval; MEDLINE","score_opus":0.04076914789316258,"score_gpt":0.3831388846104964,"score_spread":0.34236973671733384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401454847","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06455553,0.11047944,0.6627954,0.055339545,0.0017760681,0.008552908,0.020864109,0.05689334,0.018743599],"genre_scores_gemma":[0.1206138,0.01592821,0.8413927,0.0040412312,0.00067689305,0.004353791,0.010599063,0.0015281306,0.0008660906],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.7489258,0.18156716,0.03274504,0.008538509,0.026849657,0.0013738098],"domain_scores_gemma":[0.06560065,0.83071357,0.026767658,0.04688423,0.028110178,0.0019236604],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.20149796,0.0023048061,0.002873591,0.023093702,0.0015730695,0.013663672,0.0048211464,0.0039120484,0.012915548],"category_scores_gemma":[0.69599783,0.0020872075,0.004866348,0.015446385,0.0018394189,0.011281392,0.006966289,0.0027952006,0.0055350256],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001330302,0.0003291766,0.012479715,0.05232367,0.002876881,0.0003107328,0.0035184224,0.0046740994,0.003791592,0.0048443107,0.033775263,0.87974584],"study_design_scores_gemma":[0.004760798,0.002641155,0.08096423,0.12412608,0.011040727,0.003138807,0.006138474,0.14268985,0.029106209,0.1134035,0.4803935,0.0015967145],"about_ca_topic_score_codex":0.0037102692,"about_ca_topic_score_gemma":0.0072271167,"teacher_disagreement_score":0.798502,"about_ca_system_score_codex":0.0029601355,"about_ca_system_score_gemma":0.014958625,"threshold_uncertainty_score":0.9846952},"labels":[],"label_agreement":null},{"id":"W4401502666","doi":"10.1145/3659677.3659742","title":"Towards an Automatic Extracting UML Class Diagram from System's Textual Specification","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Class diagram; Computer science; Unified Modeling Language; Programming language; Applications of UML; Communication diagram; Class (philosophy); Use Case Diagram; UML tool; Natural language processing; Software engineering; Artificial intelligence; Software","score_opus":0.030457292884512897,"score_gpt":0.2978104770582071,"score_spread":0.26735318417369425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401502666","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008337024,0.00019465921,0.9794413,0.00026212377,0.000027287626,0.0001339961,0.0011254854,0.00967614,0.0008020551],"genre_scores_gemma":[0.047820408,0.00021771161,0.9451044,0.00008737002,0.000012927279,0.00014512113,0.004888846,0.0005969848,0.0011261457],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99857295,0.0003657197,0.00013407618,0.00032371536,0.0005357057,0.00006785125],"domain_scores_gemma":[0.9960343,0.001938001,0.0005409756,0.000493364,0.0008916934,0.00010163564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014205647,0.0012448867,0.00059653143,0.003362016,0.0004267321,0.0014402919,0.0010015778,0.0011926861,0.001814786],"category_scores_gemma":[0.0063075647,0.0007494247,0.0014986792,0.001251182,0.00040770406,0.0016166037,0.0010518503,0.0014577515,0.0019804046],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013168651,0.000312366,0.0055412548,0.0012900068,0.000089807625,0.0006200711,0.00069099816,0.038457606,0.10585738,0.016176097,0.021616396,0.8092163],"study_design_scores_gemma":[0.000048248177,0.00009852357,0.0024755977,0.00018017933,0.00007505063,0.00055777654,0.00024161756,0.8712572,0.06916438,0.015142176,0.0407033,0.000056038563],"about_ca_topic_score_codex":0.0053526633,"about_ca_topic_score_gemma":0.010121727,"teacher_disagreement_score":0.0053526633,"about_ca_system_score_codex":0.0010531433,"about_ca_system_score_gemma":0.0028168622,"threshold_uncertainty_score":0.010643065},"labels":[],"label_agreement":null},{"id":"W4401543865","doi":"10.1145/3643991.3644934","title":"Fine-Grained Just-In-Time Defect Prediction at the Block Level in Infrastructure-as-Code (IaC)","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Code (set theory); Block (permutation group theory); Parallel computing; Programming language; Mathematics","score_opus":0.018394784471891086,"score_gpt":0.2665918901334869,"score_spread":0.2481971056615958,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401543865","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82540417,0.0009406675,0.15758935,0.00045816286,0.000088166984,0.00020597893,0.002466911,0.011408372,0.0014382629],"genre_scores_gemma":[0.92403215,0.00016121873,0.069369614,0.00009833484,0.00002543471,0.000104292856,0.004658224,0.00017791592,0.0013728669],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998747,0.00016841617,0.0001013906,0.0005014297,0.00033355915,0.00014809011],"domain_scores_gemma":[0.9925403,0.0029196646,0.0016598055,0.0007537655,0.0017345066,0.00039202522],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016411805,0.0012789901,0.0007156593,0.0029947888,0.00029017663,0.00077220995,0.0011596793,0.0009139529,0.00044732995],"category_scores_gemma":[0.008342308,0.000318766,0.0006779475,0.0011175148,0.00039033513,0.0013616177,0.000793255,0.0011432524,0.00047392544],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004372221,0.00072058866,0.35310662,0.00044000504,0.00023393755,0.000599905,0.00047529882,0.2605569,0.014243952,0.0008827953,0.009500325,0.35880247],"study_design_scores_gemma":[0.000012429475,0.00017495843,0.02917502,0.000027925835,0.00003566316,0.00014773654,0.00007418451,0.9628174,0.0055022254,0.00078988174,0.0012181015,0.000024470359],"about_ca_topic_score_codex":0.015322908,"about_ca_topic_score_gemma":0.019696176,"teacher_disagreement_score":0.015322908,"about_ca_system_score_codex":0.00075845676,"about_ca_system_score_gemma":0.0011846895,"threshold_uncertainty_score":0.030467391},"labels":[],"label_agreement":null},{"id":"W4401635770","doi":"10.1145/3688841","title":"An Exploratory Study on Machine Learning Model Management","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Concordia University","funders":"","keywords":"Documentation; Computer science; Software versioning; Software engineering; Automation; Code refactoring; Knowledge management; Data science; Process management; Software; Engineering","score_opus":0.10743960691064752,"score_gpt":0.35192113569131234,"score_spread":0.24448152878066481,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401635770","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9524736,0.0012133108,0.023428855,0.005712843,0.000055242497,0.00065472495,0.0012213561,0.000389176,0.014850789],"genre_scores_gemma":[0.96296406,0.00091867597,0.028738348,0.00090411067,0.00004541654,0.00052831025,0.0020722277,0.00019329003,0.0036355658],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9833264,0.009811507,0.00083204714,0.0013252577,0.0037607295,0.00094407826],"domain_scores_gemma":[0.8260937,0.14105868,0.007133661,0.009813312,0.012587848,0.0033127812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022900263,0.00045578228,0.0004935883,0.0024253768,0.0021581273,0.0045162383,0.0024721152,0.001486306,0.0036587443],"category_scores_gemma":[0.13001609,0.00039861814,0.00056516955,0.0044719893,0.0016081701,0.008988138,0.0023643742,0.002994867,0.001026767],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004984131,0.0051103183,0.3319247,0.0021448836,0.00013329848,0.0017470513,0.14378522,0.008932286,0.0034727387,0.047390744,0.04158285,0.41327754],"study_design_scores_gemma":[0.00017783907,0.0021294204,0.18759534,0.0021756636,0.00014493361,0.0025243086,0.19576456,0.15414676,0.007207289,0.048059158,0.39976934,0.0003053314],"about_ca_topic_score_codex":0.0046936385,"about_ca_topic_score_gemma":0.007968163,"teacher_disagreement_score":0.022900263,"about_ca_system_score_codex":0.0030871509,"about_ca_system_score_gemma":0.003634916,"threshold_uncertainty_score":0.121109664},"labels":[],"label_agreement":null},{"id":"W4401703630","doi":"10.2139/ssrn.4927380","title":"Cognitive Biases in Natural Language: Automatically Detecting, Differentiating, and Measuring Bias in Text","year":2024,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Natural language processing; Natural (archaeology); Cognition; Natural language understanding; Natural language; Cognitive bias; Artificial intelligence; Cognitive psychology; Psychology; Geography","score_opus":0.0277453928769883,"score_gpt":0.298929141323229,"score_spread":0.2711837484462407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401703630","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.87336534,0.0016977135,0.11430773,0.00063724513,0.00021185658,0.00034630296,0.0017541883,0.002965246,0.004714344],"genre_scores_gemma":[0.9266576,0.00039107085,0.06983713,0.0002567397,0.00019761584,0.0001858895,0.0014067116,0.00022876797,0.0008384252],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958801,0.001233151,0.00045618476,0.0010623523,0.0011553225,0.00021281108],"domain_scores_gemma":[0.9369895,0.04846739,0.0067608757,0.002477448,0.004326853,0.0009780315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042114444,0.00066231453,0.0008788091,0.003676842,0.00041198084,0.0035236562,0.000846995,0.0013282367,0.002008207],"category_scores_gemma":[0.05725305,0.0003567383,0.00048916653,0.0020128924,0.0007031139,0.00441014,0.0014908918,0.0012100724,0.001415324],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022716902,0.0006516473,0.12983657,0.0015365513,0.00035624957,0.0003554029,0.0040563196,0.0021163013,0.12722337,0.00300011,0.006425474,0.72217035],"study_design_scores_gemma":[0.0005655406,0.0019401493,0.45669806,0.00056583545,0.0013064765,0.0027956606,0.005349297,0.28056338,0.16905098,0.06629071,0.014338415,0.0005355161],"about_ca_topic_score_codex":0.0013676481,"about_ca_topic_score_gemma":0.0016528709,"teacher_disagreement_score":0.0042114444,"about_ca_system_score_codex":0.0005262804,"about_ca_system_score_gemma":0.00084847194,"threshold_uncertainty_score":0.022272468},"labels":[],"label_agreement":null},{"id":"W4401752332","doi":"10.1177/17470218241280567","title":"Incremental structure building in the processing of ellipsis","year":2024,"lang":"en","type":"article","venue":"Quarterly Journal of Experimental Psychology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Ellipsis (linguistics); Antecedent (behavioral psychology); Computer science; Copying; Mechanism (biology); Parsing; Natural language processing; Artificial intelligence; Word (group theory); Linguistics; Psychology; Physics","score_opus":0.021902599627507593,"score_gpt":0.36698450395073895,"score_spread":0.3450819043232314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401752332","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9637604,0.00016926571,0.027742164,0.0002374992,0.000029606017,0.00010129947,0.0000644709,0.00077840546,0.007116882],"genre_scores_gemma":[0.97136676,0.00009336811,0.026703453,0.00010877192,0.000018791226,0.00004453564,0.00012990365,0.00023955498,0.0012948189],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.99796,0.00055225636,0.00010670664,0.0004927381,0.000720578,0.00016772748],"domain_scores_gemma":[0.98812217,0.008301608,0.0008574058,0.0016957907,0.0007771203,0.00024585737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001485902,0.00053392386,0.00060613925,0.00042227897,0.00076351524,0.0019647859,0.00089894567,0.0010583769,0.0032359418],"category_scores_gemma":[0.01503455,0.00097145303,0.0005432526,0.00034948526,0.0015911115,0.005232412,0.0016073801,0.0022566768,0.0006446375],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079617684,0.00038478404,0.007006116,0.0004890327,0.00005584801,0.0011152624,0.014248075,0.0015878794,0.8421829,0.016461916,0.0004973109,0.115174755],"study_design_scores_gemma":[0.0003755217,0.0028045136,0.1355515,0.00013415092,0.00022060827,0.0042945435,0.006385351,0.046572432,0.66317385,0.1278301,0.012263905,0.00039348536],"about_ca_topic_score_codex":0.0008314046,"about_ca_topic_score_gemma":0.00070994435,"teacher_disagreement_score":0.0032359418,"about_ca_system_score_codex":0.0004993308,"about_ca_system_score_gemma":0.0006917544,"threshold_uncertainty_score":0.010825276},"labels":[],"label_agreement":null},{"id":"W4401851884","doi":"10.53555/sfs.v10i1.2974","title":"Utilizing Machine Learning For Predicting Software Defects","year":2023,"lang":"en","type":"article","venue":"Journal of Survey in Fisheries Sciences","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Software; Artificial intelligence; Machine learning; Operating system","score_opus":0.17689606124673538,"score_gpt":0.31311648702568945,"score_spread":0.13622042577895407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401851884","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7004005,0.0023670895,0.2842522,0.000776756,0.00017651996,0.0003185668,0.0028393616,0.005157801,0.003711289],"genre_scores_gemma":[0.9226679,0.00040487922,0.07264075,0.000087923516,0.00005920456,0.00010520999,0.0031542373,0.00003388567,0.0008460325],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986687,0.00039894693,0.00013071632,0.00028355236,0.00039180115,0.00012633903],"domain_scores_gemma":[0.99407154,0.0038483546,0.00065211975,0.00033073357,0.00097537757,0.00012180105],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020503202,0.0011405379,0.0007197195,0.0053800824,0.00032384097,0.00091202214,0.00078065717,0.0011145996,0.0005323698],"category_scores_gemma":[0.008374808,0.00022079132,0.0007244304,0.0021880225,0.00021855024,0.0011726057,0.00044778554,0.0007844125,0.00048732475],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030155852,0.0010230452,0.15022868,0.00027123382,0.00026739467,0.00024562588,0.000114649774,0.2755826,0.0048440997,0.00088002376,0.0057096523,0.5605315],"study_design_scores_gemma":[0.000011787152,0.00012903752,0.00917374,0.000031223564,0.00003523883,0.00006562255,0.000037099417,0.9862903,0.0024072495,0.0010948488,0.0007072362,0.000016664135],"about_ca_topic_score_codex":0.006125263,"about_ca_topic_score_gemma":0.007477395,"teacher_disagreement_score":0.006125263,"about_ca_system_score_codex":0.0005724675,"about_ca_system_score_gemma":0.0007987277,"threshold_uncertainty_score":0.0121792555},"labels":[],"label_agreement":null},{"id":"W4401863515","doi":"10.1145/3637528.3671627","title":"GraSS: Combining Graph Neural Networks with Expert Knowledge for SAT Solver Selection","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Solver; Artificial neural network; Selection (genetic algorithm); Artificial intelligence; Graph; Knowledge graph; Boolean satisfiability problem; Machine learning; Theoretical computer science; Programming language","score_opus":0.020441748822004255,"score_gpt":0.28482377552515403,"score_spread":0.26438202670314975,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401863515","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08063603,0.0013759179,0.8888014,0.0011806418,0.00027666584,0.00037310176,0.0017453239,0.015262351,0.010348545],"genre_scores_gemma":[0.55681354,0.00043903885,0.426884,0.0010967621,0.00016970266,0.0003650018,0.006263196,0.0008602988,0.0071083917],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99909055,0.00029369266,0.00003603018,0.00027279288,0.00021169794,0.00009526334],"domain_scores_gemma":[0.99823356,0.0010361692,0.00014718952,0.00025009504,0.00025447912,0.00007844823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011617396,0.001794751,0.00079873606,0.0015776542,0.00047788827,0.00096561416,0.0025468476,0.0015918775,0.0050898646],"category_scores_gemma":[0.0061597633,0.0005384283,0.0008817561,0.0015136423,0.00059178536,0.0024001217,0.0015080615,0.0020163523,0.0012605055],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000261745,0.0002456632,0.002515037,0.00024364115,0.00013977838,0.00013660749,0.00006389221,0.63426936,0.0035160528,0.008413078,0.020230575,0.32996455],"study_design_scores_gemma":[0.000016327476,0.000023389057,0.000118335425,0.0000068054424,0.00001147101,0.000010882494,0.000009218746,0.9921112,0.00060785265,0.0063904366,0.0006899911,0.0000040583586],"about_ca_topic_score_codex":0.0079663275,"about_ca_topic_score_gemma":0.02101204,"teacher_disagreement_score":0.0079663275,"about_ca_system_score_codex":0.0013159598,"about_ca_system_score_gemma":0.0012161719,"threshold_uncertainty_score":0.017027318},"labels":[],"label_agreement":null},{"id":"W4402034432","doi":"10.32920/26871409.v1","title":"Transformer Models for Automated Bug Triaging and Duplicate Bug Detection","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; University of Manitoba; Systems, Applications & Products in Data Processing (Canada)","funders":"","keywords":"Computer science; Workflow; Software; Mean reciprocal rank; Software bug; Rank (graph theory); Data mining; Artificial intelligence; Database; Programming language","score_opus":0.03379990107120691,"score_gpt":0.3007839726597976,"score_spread":0.2669840715885907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402034432","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054127228,0.001318969,0.93637395,0.0008227561,0.00012293912,0.00013998682,0.00060035946,0.0042169383,0.0022768795],"genre_scores_gemma":[0.79012555,0.0010955101,0.19781555,0.0005492825,0.00019933165,0.00025025837,0.0028579806,0.00040112334,0.006705522],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998235,0.000629153,0.00013252995,0.0004993685,0.00035840276,0.0001455157],"domain_scores_gemma":[0.99291795,0.0036248884,0.00074610865,0.0010393991,0.001433039,0.00023858325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004264507,0.0012642768,0.0014022887,0.0025631546,0.0005439765,0.0021558406,0.0022324189,0.0014297174,0.002429276],"category_scores_gemma":[0.016158344,0.0005763488,0.0016910615,0.0022426194,0.0011109128,0.003400881,0.0019414048,0.0025003788,0.0017138068],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005253269,0.00038561816,0.006884149,0.00025411727,0.00021492827,0.00017380848,0.00028079693,0.61314446,0.003910812,0.027285397,0.0132327955,0.33370787],"study_design_scores_gemma":[0.00001104077,0.000042270265,0.00023522248,0.000004944495,0.000012791758,0.000026283176,0.000012638691,0.9883024,0.00048006623,0.010435168,0.00042880647,0.000008312088],"about_ca_topic_score_codex":0.009209722,"about_ca_topic_score_gemma":0.009038489,"teacher_disagreement_score":0.009209722,"about_ca_system_score_codex":0.0016880352,"about_ca_system_score_gemma":0.001949688,"threshold_uncertainty_score":0.022553146},"labels":[],"label_agreement":null},{"id":"W4402040489","doi":"10.1109/tse.2024.3452595","title":"RLocator: Reinforcement Learning for Bug Localization","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Waterloo","funders":"","keywords":"Computer science; Reinforcement learning; Artificial intelligence; Software bug; Programming language; Software engineering; Human–computer interaction; Machine learning; Software","score_opus":0.014284326079011332,"score_gpt":0.2518457570048798,"score_spread":0.23756143092586848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402040489","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053061582,0.0009474588,0.93197376,0.000623418,0.00010265249,0.00022154402,0.00027082293,0.010954749,0.0018440263],"genre_scores_gemma":[0.7527251,0.00028707559,0.2426705,0.00061754766,0.00009090747,0.0004996429,0.0007475874,0.00047404366,0.001887609],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977804,0.0010489271,0.00010998561,0.00048536394,0.00041559903,0.00015975471],"domain_scores_gemma":[0.989824,0.007621112,0.0006866252,0.0006490924,0.0009081605,0.00031098598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042398334,0.0016554625,0.0015700586,0.0012356055,0.0003895669,0.0008230002,0.0029956745,0.0015801821,0.0018295777],"category_scores_gemma":[0.017888134,0.0007096471,0.0008031076,0.0007386355,0.0010693438,0.0015995502,0.0016213314,0.0023970865,0.0006159026],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000249632,0.00036619502,0.0052294927,0.00018987177,0.00012865588,0.00010505915,0.000097735756,0.74808663,0.0024082947,0.003075335,0.0065800766,0.233483],"study_design_scores_gemma":[0.00002495476,0.000048813494,0.00015247449,0.0000061619917,0.000008180927,0.000012171005,0.00000422116,0.99762434,0.0005006438,0.0013969865,0.00021561056,0.0000054791885],"about_ca_topic_score_codex":0.0051592193,"about_ca_topic_score_gemma":0.006033282,"teacher_disagreement_score":0.0051592193,"about_ca_system_score_codex":0.001680993,"about_ca_system_score_gemma":0.002093189,"threshold_uncertainty_score":0.022422612},"labels":[],"label_agreement":null},{"id":"W4402100576","doi":"10.1007/978-3-031-70378-2_1","title":"VulEXplaineR: XAI for Vulnerability Detection on Assembly Code","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada); Defence Research and Development Canada; McGill University","funders":"","keywords":"Computer science; Vulnerability (computing); Code (set theory); Programming language; Computer security","score_opus":0.02503805183696138,"score_gpt":0.29059986459214643,"score_spread":0.26556181275518503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402100576","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021317024,0.000906731,0.35533577,0.00023460771,0.00026382608,0.00014583762,0.0038611297,0.6040236,0.0139115155],"genre_scores_gemma":[0.2777,0.0012349553,0.56606424,0.000702935,0.00029817896,0.00057105825,0.021228986,0.08131468,0.050884906],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991321,0.00009955063,0.00005258988,0.0002150876,0.00039733457,0.00010324652],"domain_scores_gemma":[0.9986578,0.0005672049,0.00015279847,0.0004157924,0.00014884383,0.000057597303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007233701,0.0023389661,0.00080236833,0.0019917835,0.0004568417,0.0012613974,0.002018399,0.0010527375,0.022482343],"category_scores_gemma":[0.0026919881,0.0011155196,0.0012820123,0.0011341836,0.0005110937,0.002226591,0.0016594849,0.0017385556,0.010990612],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069160154,0.00020268191,0.0048439414,0.0008509294,0.00022324966,0.00042480853,0.00036586175,0.009431654,0.049476687,0.008570997,0.22314832,0.7017693],"study_design_scores_gemma":[0.00029075635,0.0007652673,0.009922107,0.00044081034,0.00033689363,0.0018157347,0.00020396152,0.45162708,0.27805388,0.029601946,0.22662288,0.0003186574],"about_ca_topic_score_codex":0.0010651102,"about_ca_topic_score_gemma":0.0014681708,"teacher_disagreement_score":0.022482343,"about_ca_system_score_codex":0.00044787372,"about_ca_system_score_gemma":0.00059234764,"threshold_uncertainty_score":0.07521093},"labels":[],"label_agreement":null},{"id":"W4402136389","doi":"10.1007/s10664-024-10528-7","title":"Consensus task interaction trace recommender to guide developers’ software navigation","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Concordia University; Polytechnique Montréal","funders":"","keywords":"TRACE (psycholinguistics); Recommender system; Computer science; Task (project management); Software; World Wide Web; Human–computer interaction; Software engineering; Data science; Engineering; Programming language; Systems engineering","score_opus":0.036349133622559676,"score_gpt":0.32891704916676656,"score_spread":0.2925679155442069,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402136389","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22719708,0.0013764851,0.7227247,0.0013858497,0.00048140137,0.00049479207,0.002355515,0.027004274,0.016979935],"genre_scores_gemma":[0.78124505,0.0002731771,0.19527408,0.00033873584,0.00012208255,0.00027691646,0.0041349083,0.0009917214,0.017343342],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971468,0.0010019208,0.00017364146,0.00052419334,0.0009747117,0.00017875066],"domain_scores_gemma":[0.97887635,0.008946027,0.0010898497,0.0030044,0.0069270073,0.0011563293],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003049264,0.0012030613,0.0008636453,0.0039352295,0.0010446034,0.0016287094,0.0019915968,0.001932364,0.005781249],"category_scores_gemma":[0.027050672,0.0004952618,0.0006648727,0.001700822,0.0002550295,0.0036494927,0.0018392238,0.0019540212,0.0041206814],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002397458,0.0022452744,0.060153738,0.0006327002,0.00043245102,0.00036520712,0.0020687887,0.033507727,0.014506774,0.00840697,0.06715013,0.80813277],"study_design_scores_gemma":[0.00018793307,0.0009841272,0.012708109,0.00010802323,0.00024917873,0.00021768748,0.0008406983,0.9422513,0.010925959,0.0110966675,0.020287553,0.00014282136],"about_ca_topic_score_codex":0.018536326,"about_ca_topic_score_gemma":0.048473213,"teacher_disagreement_score":0.018536326,"about_ca_system_score_codex":0.0009011672,"about_ca_system_score_gemma":0.0033194688,"threshold_uncertainty_score":0.03685689},"labels":[],"label_agreement":null},{"id":"W4402315368","doi":"10.1016/j.jss.2024.112205","title":"Feature transformation for improved software bug detection and commit classification","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Commit; Transformation (genetics); Computer science; Feature (linguistics); Software bug; Software; Pattern recognition (psychology); Artificial intelligence; Data mining; Programming language; Database","score_opus":0.018545353502494814,"score_gpt":0.26088456132917104,"score_spread":0.24233920782667623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402315368","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23918508,0.00044418618,0.73745584,0.00054705876,0.00014682332,0.0001294475,0.00137969,0.019362573,0.0013493537],"genre_scores_gemma":[0.8143212,0.00010278503,0.17961109,0.000099265315,0.000040192317,0.000117548196,0.00396416,0.00032635682,0.0014173401],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985973,0.00030733712,0.00013785229,0.00034426313,0.00046035412,0.00015287548],"domain_scores_gemma":[0.9952081,0.0017008052,0.0004949332,0.0011218139,0.0013604728,0.00011389643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014684659,0.0009811845,0.0010312162,0.002665651,0.0003864727,0.0009153678,0.0012086678,0.0008467529,0.0015972507],"category_scores_gemma":[0.009477589,0.0002897517,0.0011293343,0.0024366886,0.00034959603,0.0016665292,0.0009508531,0.0017706925,0.0011151333],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002938262,0.00052236085,0.027988179,0.000109479995,0.00011426158,0.00017826662,0.00014633428,0.09881583,0.016609661,0.0020128756,0.0072459863,0.845963],"study_design_scores_gemma":[0.000020300018,0.00012890005,0.0052627237,0.000014607394,0.00003040418,0.000108837754,0.000044853718,0.9778072,0.010439564,0.004552341,0.0015668272,0.000023377572],"about_ca_topic_score_codex":0.007947538,"about_ca_topic_score_gemma":0.0064106514,"teacher_disagreement_score":0.007947538,"about_ca_system_score_codex":0.000776006,"about_ca_system_score_gemma":0.0011125692,"threshold_uncertainty_score":0.015802562},"labels":[],"label_agreement":null},{"id":"W4402400605","doi":"10.1007/978-3-031-70245-7_19","title":"Goal Model Extraction from User Stories Using Large Language Models","year":2024,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Extraction (chemistry); Natural language processing; Information retrieval; Information extraction; Artificial intelligence; Chromatography; Chemistry","score_opus":0.06199237035249826,"score_gpt":0.343000401956136,"score_spread":0.28100803160363774,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402400605","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04159257,0.0011357204,0.9135197,0.0011445733,0.00014323788,0.00049911806,0.013015269,0.017164908,0.011785012],"genre_scores_gemma":[0.2969874,0.0013422121,0.6458663,0.0003599457,0.0000893918,0.0008226159,0.04198428,0.0026107675,0.0099371355],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99881953,0.00044445746,0.00010206681,0.00023309891,0.00033809413,0.00006276435],"domain_scores_gemma":[0.9936957,0.004978153,0.00021638747,0.0003995464,0.00062977,0.000080532554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009932929,0.001796056,0.00066557137,0.0024302606,0.0007258035,0.0022048706,0.0013421586,0.0011114972,0.0065288343],"category_scores_gemma":[0.0074426956,0.0008700039,0.0016555791,0.0018027005,0.00046095665,0.0042628166,0.001680421,0.0018689585,0.0042348094],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006955823,0.00053851923,0.0068511968,0.0028789225,0.00033305818,0.0035399436,0.006465464,0.04055451,0.04080633,0.028121345,0.09051475,0.7787004],"study_design_scores_gemma":[0.00009423971,0.00020871378,0.0040658675,0.0007032244,0.00040487354,0.002511267,0.0036720857,0.71342146,0.06343422,0.047180884,0.1641278,0.0001753866],"about_ca_topic_score_codex":0.0046268185,"about_ca_topic_score_gemma":0.007719589,"teacher_disagreement_score":0.0065288343,"about_ca_system_score_codex":0.00090062,"about_ca_system_score_gemma":0.0011619123,"threshold_uncertainty_score":0.021841109},"labels":[],"label_agreement":null},{"id":"W4402400632","doi":"10.1007/s00766-024-00430-5","title":"Recommending and release planning of user-driven functionality deletion for mobile apps","year":2024,"lang":"en","type":"article","venue":"Requirements Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; York University","funders":"","keywords":"Computer science; Mobile apps; Human–computer interaction; World Wide Web","score_opus":0.0403477872092347,"score_gpt":0.3155739332268842,"score_spread":0.2752261460176495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402400632","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3216455,0.0017038895,0.6486143,0.001071332,0.00034147032,0.0014103617,0.0014388027,0.013478278,0.010296153],"genre_scores_gemma":[0.70684415,0.0003993412,0.28433588,0.00017322313,0.000049043345,0.00043433983,0.0025188187,0.00087911595,0.0043660174],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99442685,0.0016044443,0.0004248902,0.00071546505,0.002431623,0.0003967716],"domain_scores_gemma":[0.9778911,0.012326136,0.0015051306,0.0033472602,0.004144896,0.0007855412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043490813,0.0017114172,0.0009585731,0.0027401482,0.0010369408,0.0019117157,0.0017063738,0.0014521497,0.005515077],"category_scores_gemma":[0.026488295,0.0013061037,0.0019467061,0.00087103195,0.0007074158,0.0032644984,0.0014146004,0.0019041453,0.0017306752],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020653175,0.0012391456,0.036723725,0.001693872,0.00027720307,0.002788663,0.0025213973,0.13564326,0.07205969,0.008834846,0.025815811,0.7103371],"study_design_scores_gemma":[0.00013246936,0.0011065259,0.013798998,0.0003006531,0.000324547,0.0010198649,0.0014735918,0.92929775,0.033427153,0.0060282205,0.012916781,0.0001735428],"about_ca_topic_score_codex":0.0074997763,"about_ca_topic_score_gemma":0.011984439,"teacher_disagreement_score":0.0074997763,"about_ca_system_score_codex":0.00073917455,"about_ca_system_score_gemma":0.0027499525,"threshold_uncertainty_score":0.02300042},"labels":[],"label_agreement":null},{"id":"W4402475686","doi":"10.1109/ccece59415.2024.10667167","title":"Feature Importance in the Context of Traditional and Just-In-Time Software Defect Prediction Models","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Fanshawe College","funders":"","keywords":"Computer science; Context (archaeology); Feature (linguistics); Software bug; Software; Artificial intelligence; Data mining; Predictive modelling; Context model; Machine learning; Programming language","score_opus":0.03434604908915853,"score_gpt":0.2578261732150657,"score_spread":0.22348012412590715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402475686","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35703066,0.0008079466,0.6368092,0.0010478265,0.00011293778,0.00009402699,0.00037495542,0.0006024348,0.0031200687],"genre_scores_gemma":[0.9699321,0.00015144754,0.028648742,0.00006134751,0.000060989227,0.00002756905,0.00023727903,0.000022280856,0.0008582891],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99903584,0.0003484012,0.00005507778,0.00022853893,0.00022800342,0.000104288556],"domain_scores_gemma":[0.9921577,0.005581919,0.0006347259,0.0006213623,0.00080220227,0.00020197514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003247644,0.0008168303,0.0007654458,0.0014593758,0.00021757657,0.0011637338,0.0010382042,0.00096801843,0.0010051257],"category_scores_gemma":[0.011970417,0.00023981532,0.0006124612,0.0010172917,0.000768991,0.0025541831,0.0008906645,0.0016607953,0.0001484916],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028933882,0.00028934833,0.027859882,0.0001463305,0.00020520629,0.00027488352,0.00017183587,0.8203045,0.0020349512,0.027465668,0.0019400963,0.11901793],"study_design_scores_gemma":[0.000005727921,0.000042764263,0.0022645586,0.0000069280018,0.000015828666,0.000033048404,0.000010432636,0.980909,0.00051666214,0.015962962,0.00022443527,0.000007621987],"about_ca_topic_score_codex":0.0023948196,"about_ca_topic_score_gemma":0.0022745284,"teacher_disagreement_score":0.003247644,"about_ca_system_score_codex":0.0007774775,"about_ca_system_score_gemma":0.0006700155,"threshold_uncertainty_score":0.017175436},"labels":[],"label_agreement":null},{"id":"W4402483878","doi":"10.1145/3696002","title":"A Novel Refactoring and Semantic Aware Abstract Syntax Tree Differencing Tool and a Benchmark for Evaluating the Accuracy of Diff Tools","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Benchmark (surveying); Abstract syntax tree; Abstract syntax; Programming language; Syntax; Software engineering; Tree (set theory); Semantics (computer science); Artificial intelligence; Software","score_opus":0.18087197980816785,"score_gpt":0.3856275967855456,"score_spread":0.20475561697737774,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402483878","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39467457,0.005151602,0.36939597,0.00047049974,0.00057920825,0.0007837068,0.015764933,0.20786849,0.005311022],"genre_scores_gemma":[0.43415183,0.0006821627,0.51726115,0.00029977984,0.00007769156,0.0005980889,0.040732354,0.0034124015,0.0027845378],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9883265,0.0020626402,0.0021482331,0.0027654388,0.004214083,0.00048300647],"domain_scores_gemma":[0.96712923,0.014577638,0.003172623,0.0058631995,0.008528554,0.0007288902],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005836069,0.0025058198,0.0011613932,0.011168179,0.00071018527,0.0015548762,0.002982541,0.0022579846,0.0012778747],"category_scores_gemma":[0.037841298,0.0006724041,0.0011552662,0.005152533,0.0007406771,0.0033756099,0.002385216,0.0012322888,0.0012596853],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010166198,0.00071654044,0.0757986,0.0022120052,0.00044689298,0.0012675058,0.001148724,0.029288042,0.041278,0.0041210605,0.056444973,0.786261],"study_design_scores_gemma":[0.0006858154,0.0017298457,0.06468144,0.00054450077,0.00039491965,0.0040224376,0.0010063279,0.6727848,0.17367646,0.007020215,0.07286976,0.00058350037],"about_ca_topic_score_codex":0.006514304,"about_ca_topic_score_gemma":0.0069300844,"teacher_disagreement_score":0.011168179,"about_ca_system_score_codex":0.0010040193,"about_ca_system_score_gemma":0.0022347313,"threshold_uncertainty_score":0.030864477},"labels":[],"label_agreement":null},{"id":"W4402672110","doi":"10.18653/v1/2024.acl-long.361","title":"Automated Justification Production for Claim Veracity in Fact Checking: A Survey on Architectures and Approaches","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Production (economics); Computer science; Data science; Economics","score_opus":0.11318477883679201,"score_gpt":0.3217588174160254,"score_spread":0.2085740385792334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402672110","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024642345,0.033700436,0.88941914,0.00939179,0.00048145038,0.001429815,0.0016453884,0.014271973,0.025017625],"genre_scores_gemma":[0.29089162,0.014694376,0.68044245,0.0012422226,0.0005324521,0.00043885736,0.004452872,0.0017938769,0.0055112103],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.95884776,0.01676474,0.0037161561,0.005911397,0.013475259,0.001284799],"domain_scores_gemma":[0.77042985,0.15673853,0.018852683,0.031937566,0.02029702,0.0017444453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03126125,0.0022395635,0.0028142277,0.018018428,0.0034013079,0.012058476,0.006843091,0.0039200094,0.0076343515],"category_scores_gemma":[0.14450094,0.0015756359,0.003794158,0.00830639,0.0065788147,0.012816652,0.008473758,0.0044195047,0.004039432],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024579826,0.0003733975,0.0152934985,0.0034138276,0.00029285488,0.00066054997,0.0031719585,0.012294316,0.002604272,0.10898021,0.014879598,0.8377898],"study_design_scores_gemma":[0.00017365963,0.00034621442,0.016549172,0.010167446,0.00090185896,0.0044607846,0.0062036873,0.25372913,0.020214366,0.46157196,0.22511214,0.0005695732],"about_ca_topic_score_codex":0.010880388,"about_ca_topic_score_gemma":0.008203669,"teacher_disagreement_score":0.03126125,"about_ca_system_score_codex":0.004216793,"about_ca_system_score_gemma":0.010162379,"threshold_uncertainty_score":0.16532725},"labels":[],"label_agreement":null},{"id":"W4402706083","doi":"10.1145/3674805.3686685","title":"An Empirical Study of API Misuses of Data-Centric Libraries","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Application programming interface; Workflow; Documentation; World Wide Web; Software; Data science; Database; Programming language","score_opus":0.11442024848345958,"score_gpt":0.3921641979897521,"score_spread":0.27774394950629255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402706083","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9929825,0.00059106667,0.003054096,0.0004102724,0.000017912116,0.000081492595,0.00071317324,0.00039021616,0.0017594111],"genre_scores_gemma":[0.99097145,0.00067897455,0.0052571325,0.00024394275,0.000029225388,0.00012258936,0.0016156248,0.0002222957,0.0008588238],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.97826546,0.0055130064,0.003368481,0.0022499373,0.009545409,0.0010575643],"domain_scores_gemma":[0.73566747,0.14415172,0.06688772,0.018085063,0.031673945,0.0035340213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009600929,0.0006798141,0.000522977,0.0064962786,0.0011663352,0.0019291322,0.001203979,0.0011718258,0.001016063],"category_scores_gemma":[0.11162172,0.0007693871,0.0006172039,0.005268841,0.0018824609,0.007119655,0.0025190746,0.0024120365,0.0006308228],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000266717,0.00046354838,0.90945005,0.0006002135,0.0001389324,0.0011434116,0.013067148,0.00080654246,0.0028591661,0.0008507615,0.003438976,0.06691454],"study_design_scores_gemma":[0.000042928357,0.0007751953,0.90998,0.0008691771,0.0002406785,0.0060845637,0.020551821,0.021376643,0.01198944,0.0016372929,0.026283449,0.00016870227],"about_ca_topic_score_codex":0.0038552925,"about_ca_topic_score_gemma":0.005741047,"teacher_disagreement_score":0.009600929,"about_ca_system_score_codex":0.0010434352,"about_ca_system_score_gemma":0.0016696706,"threshold_uncertainty_score":0.05077511},"labels":[],"label_agreement":null},{"id":"W4402978017","doi":"10.1145/3640310.3674081","title":"Automated Derivation of UML Sequence Diagrams from User Stories: Unleashing the Power of Generative AI vs. a Rule-Based Approach","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; University of Prince Edward Island; Thompson Rivers University","funders":"","keywords":"Sequence diagram; Unified Modeling Language; Computer science; Sequence (biology); Programming language; Class diagram; Generative grammar; Activity diagram; Artificial intelligence; Natural language processing; Software engineering; Software","score_opus":0.028708444157130784,"score_gpt":0.2890107470865002,"score_spread":0.2603023029293694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402978017","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0699503,0.00044948392,0.9140021,0.0005844903,0.000049238522,0.0005755686,0.0020741385,0.009596928,0.0027177099],"genre_scores_gemma":[0.17823566,0.00029238002,0.8105911,0.00012484759,0.000015852378,0.00038981996,0.008434059,0.0008515803,0.0010646327],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9936085,0.003782027,0.00040977335,0.0006319045,0.0014529515,0.00011476917],"domain_scores_gemma":[0.9523922,0.03514046,0.0015505957,0.006642237,0.0039393664,0.00033499402],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050395,0.0009028006,0.00059332204,0.0032390761,0.0007399801,0.0018971665,0.0013615628,0.0011118234,0.001719969],"category_scores_gemma":[0.038835358,0.0006728048,0.0010615536,0.0016308755,0.0006849588,0.002184364,0.0018326462,0.0013088599,0.0009625302],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003363547,0.00049777795,0.010485254,0.0015113057,0.0001793221,0.0012154753,0.0038771976,0.082526155,0.02964572,0.02082384,0.011134863,0.83776677],"study_design_scores_gemma":[0.0001239849,0.0001775458,0.003856432,0.00032563086,0.000078592675,0.00082315697,0.0010270968,0.88192314,0.04291499,0.028367877,0.040307466,0.000074132324],"about_ca_topic_score_codex":0.0053367037,"about_ca_topic_score_gemma":0.010865465,"teacher_disagreement_score":0.0053367037,"about_ca_system_score_codex":0.0010387963,"about_ca_system_score_gemma":0.0023155988,"threshold_uncertainty_score":0.02665174},"labels":[],"label_agreement":null},{"id":"W4403155837","doi":"10.1007/s10115-024-02248-7","title":"Deep-transfer learning inspired natural language processing system for software requirements classification","year":2024,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Artificial intelligence; Transfer of learning; Natural language processing; Natural language; Deep learning; Software; Machine learning; Programming language","score_opus":0.0214657515274903,"score_gpt":0.28490399424704965,"score_spread":0.26343824271955935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403155837","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09917299,0.0003920372,0.8707159,0.00057623425,0.00012665916,0.00024126169,0.001673524,0.020002393,0.007098961],"genre_scores_gemma":[0.61114985,0.000267145,0.37096277,0.00055475073,0.00005719994,0.00028859475,0.00441062,0.0002669018,0.012042177],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997296,0.000045208155,0.000020439678,0.00008203543,0.00007750637,0.000045255216],"domain_scores_gemma":[0.9995628,0.00015526358,0.000038942966,0.00006622779,0.00014457948,0.000032242944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004312079,0.0004911506,0.00041633451,0.00080115284,0.00033616603,0.0005081446,0.0009636979,0.0006290975,0.004044031],"category_scores_gemma":[0.0010545794,0.00019131179,0.00060348934,0.0006454354,0.00019345172,0.001064277,0.0007184073,0.0012914612,0.001171231],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000227827,0.00072550215,0.002797851,0.00020487804,0.00008437915,0.00025882543,0.00015972329,0.06259241,0.027754977,0.006419834,0.021601856,0.877172],"study_design_scores_gemma":[0.000013400876,0.000065156346,0.00055428257,0.0000092338405,0.000015541302,0.000047077476,0.00002765542,0.98500234,0.0074311104,0.004600511,0.0022251043,0.0000086719],"about_ca_topic_score_codex":0.0084959855,"about_ca_topic_score_gemma":0.013670697,"teacher_disagreement_score":0.0084959855,"about_ca_system_score_codex":0.00089629943,"about_ca_system_score_gemma":0.0015135717,"threshold_uncertainty_score":0.016893089},"labels":[],"label_agreement":null},{"id":"W4403210207","doi":"10.1109/access.2024.3473942","title":"Ensemble Balanced Nested Dichotomy Fuzzy Models for Software Requirement Risk Prediction","year":2024,"lang":"en","type":"article","venue":"IEEE Access","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Universiti Teknologi Malaysia; Universiti Teknologi Petronas; Ministry of Higher Education and Scientific Research","keywords":"Computer science; Nested set model; Fuzzy logic; Software; Fuzzy set; Data mining; Artificial intelligence; Programming language","score_opus":0.046387383863068024,"score_gpt":0.3226909148801467,"score_spread":0.27630353101707866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403210207","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17997412,0.0010002715,0.81460416,0.0005120232,0.000062871164,0.00008528464,0.00038847723,0.00033866367,0.0030341337],"genre_scores_gemma":[0.9386188,0.00031724665,0.059049603,0.000083270876,0.00003528827,0.00009741092,0.0004195543,0.000014387678,0.0013644049],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991943,0.00023961696,0.000049469454,0.00018324213,0.0002467434,0.00008663979],"domain_scores_gemma":[0.9980038,0.0012617086,0.00019719901,0.00010467229,0.00037094834,0.000061631865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001902733,0.0006846736,0.00093250314,0.0011531427,0.0004893969,0.0007885735,0.0012446963,0.0008017534,0.0012102051],"category_scores_gemma":[0.004275628,0.00026593395,0.0010015606,0.0006787032,0.00032272283,0.0009873065,0.0006986211,0.0013697755,0.00022732238],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021078986,0.00013027656,0.008788805,0.00006823333,0.00011647174,0.000114657596,0.00019490693,0.8686009,0.0015511478,0.009172153,0.0009956743,0.110055886],"study_design_scores_gemma":[0.0000023600064,0.000015219521,0.00045180973,0.0000057705497,0.000007860569,0.0000066867897,0.000008875606,0.9969549,0.00012453856,0.0023053922,0.00011294176,0.0000035836517],"about_ca_topic_score_codex":0.010252685,"about_ca_topic_score_gemma":0.009972246,"teacher_disagreement_score":0.010252685,"about_ca_system_score_codex":0.0009385648,"about_ca_system_score_gemma":0.0007943484,"threshold_uncertainty_score":0.02038604},"labels":[],"label_agreement":null},{"id":"W4403255280","doi":"10.1007/s10664-024-10441-z","title":"Evaluating few-shot and contrastive learning methods for code clone detection","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; clone (Java method); Artificial intelligence; Code (set theory); Shot (pellet); Programming language; Biology; DNA; Genetics","score_opus":0.09238815332234587,"score_gpt":0.4441704756515731,"score_spread":0.35178232232922724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403255280","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6102431,0.021572413,0.33602673,0.0028277491,0.0011172328,0.000708886,0.0019013962,0.01667543,0.008927023],"genre_scores_gemma":[0.8476381,0.0011025143,0.13945791,0.0012943097,0.0002856198,0.00027300522,0.004502909,0.0006320029,0.0048135994],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99585366,0.0013684031,0.0001940167,0.001503742,0.00080778723,0.00027252876],"domain_scores_gemma":[0.9846585,0.011092508,0.0008159975,0.0012711341,0.0013924132,0.00076949474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007512184,0.0027867137,0.0019131626,0.0022082124,0.0008458235,0.0019495134,0.004643397,0.004193611,0.0015050604],"category_scores_gemma":[0.023489539,0.0006978136,0.0013792063,0.0010675906,0.0019175475,0.0043035233,0.0026778856,0.0044213724,0.0008420239],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019957516,0.0022085437,0.014532705,0.0011966467,0.00076595903,0.00025764472,0.00035844988,0.5333946,0.010001791,0.00373464,0.0135085685,0.41804475],"study_design_scores_gemma":[0.00004312371,0.00029439986,0.0009144303,0.00004008291,0.00003972303,0.000057326164,0.000038808063,0.9929969,0.0033753563,0.0015145045,0.000666724,0.000018792573],"about_ca_topic_score_codex":0.009342107,"about_ca_topic_score_gemma":0.010197704,"teacher_disagreement_score":0.009342107,"about_ca_system_score_codex":0.0031022925,"about_ca_system_score_gemma":0.001626989,"threshold_uncertainty_score":0.0397287},"labels":[],"label_agreement":null},{"id":"W4403413321","doi":"10.1145/3674805.3686689","title":"Are Large Language Models a Threat to Programming Platforms? An Exploratory Study","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Programming language; Exploratory research","score_opus":0.06051735923651911,"score_gpt":0.3359767694065423,"score_spread":0.2754594101700232,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403413321","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9969207,0.00004837787,0.0007407622,0.00029895792,0.0000047455483,0.00006309639,0.000023752125,0.000011323441,0.0018883016],"genre_scores_gemma":[0.9979267,0.000066316985,0.0011675241,0.00013166275,0.000008060023,0.00012439249,0.000041442327,0.000019192867,0.00051463046],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.98390156,0.009962335,0.00067636,0.0011984005,0.0031494498,0.0011117735],"domain_scores_gemma":[0.8909877,0.07695736,0.015925404,0.005491969,0.0061252215,0.00451232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01550873,0.0005422155,0.00050569966,0.0020998078,0.002475305,0.0050525093,0.0015263109,0.0011366041,0.00260348],"category_scores_gemma":[0.08287891,0.0006560073,0.00040138577,0.0012069714,0.004113327,0.0072820866,0.0046534287,0.0027079324,0.00058548053],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038128142,0.0021066326,0.41415754,0.00053032546,0.00006368456,0.0017762484,0.5218565,0.0005807403,0.0031811642,0.0050403913,0.002359512,0.04796601],"study_design_scores_gemma":[0.000053722044,0.0019218324,0.2804319,0.00055795035,0.00005800114,0.0024487572,0.6808575,0.0058765938,0.0021006293,0.0045350078,0.021033645,0.00012449152],"about_ca_topic_score_codex":0.0015499989,"about_ca_topic_score_gemma":0.0025652216,"teacher_disagreement_score":0.01550873,"about_ca_system_score_codex":0.0015113105,"about_ca_system_score_gemma":0.0018604605,"threshold_uncertainty_score":0.08201897},"labels":[],"label_agreement":null},{"id":"W4403413484","doi":"10.1145/3674805.3686677","title":"Reevaluating the Defect Proneness of Atoms of Confusion in Java Systems","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Confusion; Java; Computer science; Programming language; Psychology","score_opus":0.028538223980995647,"score_gpt":0.3039453469262527,"score_spread":0.27540712294525704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403413484","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9888033,0.00016593086,0.008333064,0.00047136642,0.000011132536,0.000056643185,0.0000638591,0.00018163213,0.0019130637],"genre_scores_gemma":[0.99434906,0.000031457505,0.005300135,0.000051707848,0.000006005412,0.000013551423,0.0000550832,0.000033974247,0.0001589691],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9880155,0.0036463707,0.0010266872,0.0013857562,0.0054424396,0.00048334236],"domain_scores_gemma":[0.63768226,0.24755947,0.056613613,0.016989714,0.03775021,0.0034047326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015991984,0.0005425631,0.00033384893,0.0039807856,0.00082444405,0.0027687068,0.0013744517,0.0012282383,0.0011556112],"category_scores_gemma":[0.1679917,0.00046080965,0.0004968901,0.0015358164,0.0025900023,0.0058007333,0.003104215,0.0017058809,0.00019831312],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010455113,0.00043274675,0.8310435,0.00044327424,0.0002202074,0.00090471894,0.016598584,0.016668877,0.0071846596,0.0068667973,0.0011075643,0.11748358],"study_design_scores_gemma":[0.000120631055,0.0022698224,0.73098105,0.00054921873,0.00052545563,0.0026983074,0.011685438,0.20832402,0.016076945,0.022500178,0.004027493,0.000241387],"about_ca_topic_score_codex":0.0055974103,"about_ca_topic_score_gemma":0.0058404263,"teacher_disagreement_score":0.015991984,"about_ca_system_score_codex":0.0022738522,"about_ca_system_score_gemma":0.0019327501,"threshold_uncertainty_score":0.0845747},"labels":[],"label_agreement":null},{"id":"W4403536026","doi":"10.1145/3691620.3695555","title":"VulAdvisor: Natural Language Suggestion Generation for Software Vulnerability Repair","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Research Foundation Singapore","keywords":"Computer science; Vulnerability (computing); Natural language generation; Natural language; Computer security; Natural language processing","score_opus":0.02011497258553036,"score_gpt":0.3067431268636739,"score_spread":0.28662815427814353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403536026","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057379395,0.0021084205,0.5015277,0.0011893029,0.00063804537,0.0012658047,0.03484997,0.39514136,0.0058999714],"genre_scores_gemma":[0.14420675,0.00054047373,0.7665783,0.00085734006,0.00012970294,0.0017239361,0.07426962,0.00515151,0.00654234],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99767905,0.0008353964,0.00014095662,0.0007791806,0.0004703566,0.00009509592],"domain_scores_gemma":[0.9939944,0.00391074,0.0002670233,0.001052461,0.00061134226,0.0001641134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017438807,0.0026533057,0.0007102296,0.0023594878,0.0005516218,0.0009187716,0.0028368346,0.0018647725,0.008506588],"category_scores_gemma":[0.011872735,0.00057924277,0.0011419952,0.0011230663,0.00058715645,0.0023985647,0.00224801,0.0025280938,0.005194533],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009238472,0.00071498996,0.0066791344,0.0029017832,0.00020694477,0.0009964061,0.0013673544,0.024791325,0.033936158,0.005332333,0.25736213,0.66478765],"study_design_scores_gemma":[0.0005080143,0.00064335,0.005397181,0.00026879864,0.00015630334,0.0011800793,0.00065473485,0.74313414,0.055235777,0.018787988,0.17382264,0.00021097281],"about_ca_topic_score_codex":0.0036722124,"about_ca_topic_score_gemma":0.011751505,"teacher_disagreement_score":0.008506588,"about_ca_system_score_codex":0.0007873001,"about_ca_system_score_gemma":0.0015972188,"threshold_uncertainty_score":0.028457403},"labels":[],"label_agreement":null},{"id":"W4403536501","doi":"10.1145/3691620.3695514","title":"Understanding the Implications of Changes to Build Systems","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Waterloo","funders":"","keywords":"Computer science","score_opus":0.1282729900227623,"score_gpt":0.32587025571332223,"score_spread":0.19759726569055994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403536501","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4411294,0.0123557355,0.060893293,0.123552315,0.0012082794,0.00028317902,0.0010143269,0.0005852508,0.35897833],"genre_scores_gemma":[0.98154336,0.0032957175,0.0051737013,0.00250051,0.00026134477,0.00005688852,0.00025152508,0.00011935262,0.0067974795],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9947514,0.0019333925,0.00020220209,0.00077909103,0.0014635365,0.00087044906],"domain_scores_gemma":[0.95082897,0.032891322,0.005256758,0.0035327794,0.006035038,0.001455117],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043405187,0.0004776608,0.00026572513,0.001752959,0.0018643618,0.0059405905,0.0021619906,0.0032269552,0.014229266],"category_scores_gemma":[0.056702305,0.0004886143,0.00042975828,0.0018709048,0.005091442,0.016242612,0.0029065968,0.0035631862,0.0014773225],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044500123,0.00041150002,0.047195517,0.0008846768,0.00006909139,0.002418589,0.011256515,0.022712689,0.0036792452,0.6857197,0.022139229,0.20306838],"study_design_scores_gemma":[0.00005954305,0.00017257458,0.048946887,0.00043458206,0.00008006489,0.0006889264,0.01578123,0.018136328,0.002559916,0.8196453,0.09340546,0.00008927098],"about_ca_topic_score_codex":0.009429905,"about_ca_topic_score_gemma":0.010123266,"teacher_disagreement_score":0.014229266,"about_ca_system_score_codex":0.0039873933,"about_ca_system_score_gemma":0.0028189083,"threshold_uncertainty_score":0.04760164},"labels":[],"label_agreement":null},{"id":"W4403539235","doi":"10.1007/s10664-024-10570-5","title":"Bridging the language gap: an empirical study of bindings for open source machine learning libraries across software package ecosystems","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Bridging (networking); Computer science; Open source software; Open source; Empirical research; Software engineering; Software; Programming language","score_opus":0.03953274196089684,"score_gpt":0.34172572270005686,"score_spread":0.30219298073916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403539235","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99099064,0.0001073221,0.002613867,0.00049615913,0.0000065589634,0.000020433477,0.00005483139,0.00007556528,0.005634565],"genre_scores_gemma":[0.9968849,0.00004410376,0.0012997296,0.00015926115,0.0000067557494,0.000021521839,0.00008382611,0.00008816029,0.0014118438],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98867756,0.0057535004,0.00071128993,0.000991903,0.002715489,0.001150269],"domain_scores_gemma":[0.81971127,0.1225405,0.023770226,0.012045997,0.016481733,0.0054502715],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.013747537,0.00020640437,0.00043181443,0.002911683,0.003604331,0.0047553806,0.0020429292,0.001853136,0.00652807],"category_scores_gemma":[0.15339105,0.0005185639,0.0003015993,0.004574368,0.0053493087,0.015088783,0.0088381395,0.0032614213,0.0012894161],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010960987,0.002438248,0.6631146,0.0003778761,0.00009886343,0.0008640021,0.16301115,0.0014129771,0.006679662,0.046857078,0.0036470655,0.11040237],"study_design_scores_gemma":[0.0001404381,0.0011489317,0.6161222,0.00070891104,0.00023654132,0.0015530214,0.2644314,0.019641196,0.007795387,0.05528515,0.03271612,0.00022069871],"about_ca_topic_score_codex":0.011126456,"about_ca_topic_score_gemma":0.012072071,"teacher_disagreement_score":0.99795705,"about_ca_system_score_codex":0.0018670629,"about_ca_system_score_gemma":0.0036788075,"threshold_uncertainty_score":0.07270485},"labels":[],"label_agreement":null},{"id":"W4403600269","doi":"10.1016/j.jss.2024.112266","title":"Refining software defect prediction through attentive neural models for code understanding","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Refining (metallurgy); Computer science; Code (set theory); Software; Artificial intelligence; Artificial neural network; Software bug; Machine learning; Software engineering; Programming language; Chemistry","score_opus":0.08898464247298289,"score_gpt":0.2992358029667185,"score_spread":0.2102511604937356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403600269","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43620017,0.0008139951,0.55300283,0.0008857347,0.00014309873,0.00010197695,0.00043924473,0.0038991699,0.0045138407],"genre_scores_gemma":[0.95820165,0.00013034248,0.03990252,0.00012770214,0.00002914395,0.000030279723,0.00028452507,0.00007190498,0.0012218672],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997123,0.000055605968,0.000016103928,0.000121712124,0.00005048786,0.000043721524],"domain_scores_gemma":[0.99712306,0.0019010678,0.00023333346,0.00023963297,0.00042840734,0.00007451515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000696017,0.0008803378,0.0005719523,0.0008677548,0.00025801256,0.0010398226,0.0016075309,0.0012086873,0.0022200001],"category_scores_gemma":[0.0056285174,0.0005170499,0.00069342233,0.0004160128,0.0004154306,0.002472209,0.00081886776,0.0022220016,0.00047814316],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049811864,0.0008006462,0.012333821,0.0001587984,0.00016867947,0.00023791834,0.00021897466,0.6233373,0.016062615,0.004569876,0.0038181213,0.3377951],"study_design_scores_gemma":[0.000003176789,0.000020054771,0.00038284127,0.0000044096473,0.000009358596,0.0000070886285,0.0000059687736,0.9967018,0.0008876111,0.0019177347,0.000056913363,0.000003033002],"about_ca_topic_score_codex":0.010627939,"about_ca_topic_score_gemma":0.01758582,"teacher_disagreement_score":0.010627939,"about_ca_system_score_codex":0.00093769596,"about_ca_system_score_gemma":0.0008586666,"threshold_uncertainty_score":0.021132112},"labels":[],"label_agreement":null},{"id":"W4403604523","doi":"10.1145/3672448","title":"Understanding Test Convention Consistency as a Dimension of Test Quality","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Test (biology); Consistency (knowledge bases); Dimension (graph theory); Quality (philosophy); Reliability engineering; Artificial intelligence; Mathematics; Engineering","score_opus":0.22788467896783968,"score_gpt":0.37761842431744,"score_spread":0.14973374534960035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403604523","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5015908,0.0015886242,0.4844935,0.0030794314,0.0000860504,0.00021323816,0.00044403205,0.0012128347,0.0072915936],"genre_scores_gemma":[0.90789664,0.00019569238,0.09062272,0.00024375385,0.000060348382,0.00013575047,0.00038326127,0.00018619838,0.00027555393],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.95267165,0.019985873,0.005214783,0.004211188,0.016191337,0.0017252881],"domain_scores_gemma":[0.61564565,0.27121887,0.039628122,0.03777708,0.032905385,0.0028248471],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03141218,0.0009803295,0.0010024451,0.0069274693,0.00086744194,0.0070457985,0.0018721287,0.0016834674,0.0009837581],"category_scores_gemma":[0.23781149,0.0007444669,0.0011029274,0.0054456014,0.0046758913,0.013153047,0.0033192888,0.0035672027,0.0001471649],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005151683,0.00044878488,0.5132155,0.000671176,0.0004798431,0.0004797493,0.008555736,0.05938972,0.017570622,0.14265892,0.0023051433,0.25370958],"study_design_scores_gemma":[0.00018881826,0.0017980611,0.32766014,0.00081040943,0.00046213713,0.0016116454,0.005584556,0.3100977,0.0244578,0.31143975,0.015471634,0.00041734008],"about_ca_topic_score_codex":0.0048261248,"about_ca_topic_score_gemma":0.0029193927,"teacher_disagreement_score":0.9685878,"about_ca_system_score_codex":0.0024480997,"about_ca_system_score_gemma":0.0031764111,"threshold_uncertainty_score":0.16612548},"labels":[],"label_agreement":null},{"id":"W4403606370","doi":"10.5753/sbes.2024.3348","title":"Explorando a detecção de conflitos semânticos nas integrações de código em múltiplos métodos","year":2024,"lang":"pt","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Physics","score_opus":0.039884342222816366,"score_gpt":0.3074434803395143,"score_spread":0.26755913811669796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403606370","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25981095,0.0030301376,0.7105023,0.001747522,0.00022527126,0.0008892444,0.0008203721,0.0089154765,0.014058784],"genre_scores_gemma":[0.502252,0.0012144776,0.48574784,0.00037275194,0.000044290868,0.00055402645,0.001407363,0.0018553457,0.0065519367],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9891262,0.0023070304,0.0011620919,0.0016648674,0.005012231,0.00072755775],"domain_scores_gemma":[0.9742691,0.01175701,0.002070425,0.005907648,0.0053051044,0.00069066376],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009175989,0.0016337506,0.0011909854,0.0031722153,0.0017109326,0.007198741,0.0026009898,0.0019283014,0.004214737],"category_scores_gemma":[0.037277184,0.0014941707,0.0019115972,0.002963746,0.002228679,0.01154117,0.007076837,0.003271295,0.0013416039],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00152984,0.00072652026,0.062248383,0.00410798,0.0005076526,0.0011573158,0.015527892,0.027748693,0.10478851,0.03987227,0.0069588837,0.734826],"study_design_scores_gemma":[0.00028943192,0.0021823759,0.05249554,0.0019670965,0.0012951054,0.0030581094,0.019669436,0.3284027,0.2538364,0.09330418,0.2429186,0.0005810543],"about_ca_topic_score_codex":0.007834125,"about_ca_topic_score_gemma":0.011887317,"teacher_disagreement_score":0.009175989,"about_ca_system_score_codex":0.0021493747,"about_ca_system_score_gemma":0.004314063,"threshold_uncertainty_score":0.048527837},"labels":[],"label_agreement":null},{"id":"W4403619954","doi":"10.1007/978-3-031-75764-8_9","title":"Adversarial Analysis of Software Composition Analysis Tools","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Composition (language); Adversarial system; Software; Software engineering; Programming language; Artificial intelligence","score_opus":0.017081283697743654,"score_gpt":0.2664037143505869,"score_spread":0.24932243065284326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403619954","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008285964,0.00015211699,0.9860042,0.0002648804,0.000036909794,0.00003587434,0.000036545523,0.0006835633,0.0044999784],"genre_scores_gemma":[0.6293469,0.00052360335,0.34754148,0.00028290317,0.00020319708,0.0001955436,0.00034946122,0.0006970421,0.020859985],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9961928,0.0013622346,0.00010775458,0.0004147471,0.0016053848,0.00031718772],"domain_scores_gemma":[0.9897921,0.0074160066,0.0004632609,0.0014609279,0.00068836473,0.00017939661],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031989813,0.00101438,0.00076756335,0.0012874459,0.0006561877,0.0014593182,0.0017659595,0.0010190333,0.005313566],"category_scores_gemma":[0.0130847925,0.00058227894,0.001087534,0.00092871656,0.0021289617,0.0023082045,0.0030377312,0.0026213299,0.0010756805],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022998704,0.000092251656,0.0008149939,0.00011320191,0.000089912406,0.0001150752,0.0001313804,0.54063517,0.006564521,0.28525048,0.00528405,0.16067898],"study_design_scores_gemma":[0.0000048238576,0.000020866559,0.00007987497,0.000013106678,0.000008479806,0.00003115026,0.000007629464,0.9119357,0.0019076619,0.08476642,0.0012188827,0.000005488226],"about_ca_topic_score_codex":0.0009937709,"about_ca_topic_score_gemma":0.0011912805,"teacher_disagreement_score":0.005313566,"about_ca_system_score_codex":0.0015372033,"about_ca_system_score_gemma":0.0010939523,"threshold_uncertainty_score":0.017775655},"labels":[],"label_agreement":null},{"id":"W4403633511","doi":"10.1007/978-3-031-70285-3_4","title":"Cascade Generalization-Based Classifiers for Software Defect Prediction","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Cascade; Computer science; Generalization; Artificial intelligence; Software; Machine learning; Pattern recognition (psychology); Data mining; Programming language; Mathematics; Engineering","score_opus":0.019014011630154867,"score_gpt":0.24246077395973276,"score_spread":0.22344676232957789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403633511","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.099551916,0.005571034,0.8844335,0.00036037178,0.00042018754,0.00016228016,0.0008029479,0.0033236067,0.0053741196],"genre_scores_gemma":[0.81055343,0.0016651956,0.17189619,0.00018000789,0.00040504173,0.00023893247,0.0020990286,0.00015531272,0.012806807],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904853,0.00020627472,0.00006849733,0.00026395594,0.0003048377,0.000107917025],"domain_scores_gemma":[0.9980191,0.0010186619,0.00011011647,0.0002621775,0.000535263,0.00005459802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001997105,0.0011248147,0.0020093226,0.0015344956,0.00054146274,0.0007203834,0.0019246761,0.0014483078,0.0027129154],"category_scores_gemma":[0.002798447,0.00046884175,0.0011118446,0.001442581,0.00030861452,0.0016100415,0.00097459473,0.0016049907,0.0014608449],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048261802,0.0003160969,0.002805161,0.0001327625,0.00018298559,0.00010386369,0.000055976532,0.14744435,0.008405117,0.0026388492,0.012497196,0.8249351],"study_design_scores_gemma":[0.000005996111,0.000073209325,0.0009493557,0.000009207185,0.000031745116,0.0000351765,0.000006298618,0.9948021,0.0017776474,0.0016651998,0.0006359793,0.000008027249],"about_ca_topic_score_codex":0.0043977886,"about_ca_topic_score_gemma":0.006029892,"teacher_disagreement_score":0.0043977886,"about_ca_system_score_codex":0.0007669211,"about_ca_system_score_gemma":0.0006648473,"threshold_uncertainty_score":0.010561764},"labels":[],"label_agreement":null},{"id":"W4403651448","doi":"10.1007/s10664-024-10568-z","title":"The downside of functional constructs: a quantitative and qualitative analysis of their fix-inducing effects","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Ministero dell'Università e della Ricerca","keywords":"Qualitative analysis; Quantitative analysis (chemistry); Downside risk; Computer science; Qualitative research; Chemistry; Economics; Sociology; Chromatography; Financial economics","score_opus":0.03409354701278354,"score_gpt":0.3381684211001117,"score_spread":0.30407487408732814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403651448","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9411356,0.0003783834,0.04537841,0.00035472008,0.000020868205,0.0001424596,0.00020096578,0.00017426522,0.012214313],"genre_scores_gemma":[0.98708296,0.000116384996,0.011805455,0.00005348787,0.000007737994,0.000085710664,0.00008478209,0.00009898448,0.00066445593],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.99496645,0.0022391044,0.0002775899,0.0004392447,0.0018227493,0.00025480028],"domain_scores_gemma":[0.8143054,0.16177931,0.0061880173,0.009640874,0.007254903,0.0008314608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008506766,0.00034263643,0.00031012818,0.002076471,0.0008498771,0.0009949224,0.00054264534,0.0005803342,0.0049135825],"category_scores_gemma":[0.06689005,0.0002952714,0.00043015214,0.0010638812,0.0019429862,0.0019041001,0.0013047602,0.0015700803,0.00021344953],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038209509,0.0012728883,0.17209977,0.0035617894,0.0003273404,0.0009745631,0.02177597,0.010141592,0.28158695,0.13603576,0.0018068233,0.36659554],"study_design_scores_gemma":[0.00022847038,0.004021725,0.63032883,0.0009570028,0.0012681886,0.0022545746,0.01652616,0.026363581,0.22937003,0.06910177,0.019364072,0.00021569383],"about_ca_topic_score_codex":0.00061764993,"about_ca_topic_score_gemma":0.0009398369,"teacher_disagreement_score":0.008506766,"about_ca_system_score_codex":0.0007392488,"about_ca_system_score_gemma":0.0007693853,"threshold_uncertainty_score":0.044988573},"labels":[],"label_agreement":null},{"id":"W4403785165","doi":"10.2139/ssrn.5000516","title":"Bugmentor: Generating Answers to Follow-Up Questions from Software Bug Reports Using Structured Information Retrieval and Neural Text Generation","year":2024,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Information retrieval; Natural language processing; Software; Artificial intelligence; Programming language","score_opus":0.0154761141132388,"score_gpt":0.2686534418935989,"score_spread":0.2531773277803601,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403785165","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08720914,0.00089742045,0.6690516,0.0010491328,0.0005006505,0.0016156404,0.013860932,0.22033718,0.005478381],"genre_scores_gemma":[0.22906128,0.00032360162,0.7278938,0.00035194142,0.00017104107,0.0010963382,0.02869756,0.0030725712,0.009331756],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985607,0.0005054885,0.000103519844,0.00037307237,0.00038085313,0.00007638419],"domain_scores_gemma":[0.9951781,0.0032614714,0.0003106897,0.00040800616,0.00068189396,0.00015983298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016567664,0.002224535,0.0011166141,0.0033811857,0.000531184,0.0010650798,0.0022507932,0.0022087346,0.016945785],"category_scores_gemma":[0.010450588,0.0006025264,0.0011618835,0.0013252251,0.00042204736,0.0019797385,0.0016205605,0.0010973666,0.006288626],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011607419,0.00081936433,0.0052048857,0.0016571761,0.00025541507,0.0008021365,0.00082081597,0.016659461,0.039519496,0.0032614733,0.091176085,0.83866286],"study_design_scores_gemma":[0.00067716837,0.00095837354,0.004812908,0.00014781488,0.00024013278,0.00063000375,0.0005527056,0.8904791,0.0625692,0.010712991,0.028090317,0.00012920707],"about_ca_topic_score_codex":0.0035856923,"about_ca_topic_score_gemma":0.005308551,"teacher_disagreement_score":0.016945785,"about_ca_system_score_codex":0.00068763946,"about_ca_system_score_gemma":0.0010204697,"threshold_uncertainty_score":0.056689262},"labels":[],"label_agreement":null},{"id":"W4403936231","doi":"10.1016/j.jss.2024.112263","title":"Impact of methodological choices on the analysis of code metrics and maintenance","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; University of British Columbia","funders":"","keywords":"Code (set theory); Computer science; Reliability engineering; Risk analysis (engineering); Engineering; Programming language; Business","score_opus":0.09372169706085566,"score_gpt":0.36956596050627394,"score_spread":0.2758442634454183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403936231","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6102023,0.010611418,0.33588094,0.020660058,0.0012687657,0.001159493,0.0017795516,0.00077081705,0.017666666],"genre_scores_gemma":[0.78385127,0.0011576043,0.21060397,0.0013543468,0.0002775945,0.0006082228,0.00052723405,0.00069358887,0.00092614035],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.582711,0.30668402,0.04519888,0.013690299,0.048881397,0.002834443],"domain_scores_gemma":[0.09527343,0.7914776,0.038119756,0.03658053,0.03680748,0.0017411542],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.30441746,0.0012128113,0.0012337855,0.009921821,0.0024041599,0.008662905,0.0039471462,0.0032594923,0.0021156783],"category_scores_gemma":[0.7430866,0.0010645478,0.0025174469,0.010642154,0.003835124,0.0069675986,0.004761421,0.003734008,0.00042675002],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010343547,0.0020221758,0.26945156,0.0061365655,0.006762833,0.0008785351,0.012086052,0.031729694,0.021058057,0.14157045,0.00731042,0.49065012],"study_design_scores_gemma":[0.003830013,0.008477894,0.26947036,0.009266078,0.009513039,0.0022730445,0.01400317,0.20493586,0.07062397,0.3625145,0.043716528,0.0013755836],"about_ca_topic_score_codex":0.004835386,"about_ca_topic_score_gemma":0.008247144,"teacher_disagreement_score":0.6955825,"about_ca_system_score_codex":0.0063930955,"about_ca_system_score_gemma":0.008796841,"threshold_uncertainty_score":0.8577771},"labels":[],"label_agreement":null},{"id":"W4403946944","doi":"10.1007/s10664-024-10560-7","title":"A qualitative study on refactorings induced by code review","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Code (set theory); Programming language; Qualitative research; Software engineering; Sociology","score_opus":0.0852812258416174,"score_gpt":0.41490363149144455,"score_spread":0.32962240564982714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403946944","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95841634,0.0006811531,0.015602919,0.0045101596,0.00023100372,0.0009860046,0.00070021994,0.00016488055,0.018707262],"genre_scores_gemma":[0.98478216,0.00050580705,0.005477158,0.0017988152,0.00006060658,0.0008451192,0.0003004193,0.00014612853,0.0060838074],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9702351,0.021244284,0.0010053546,0.0015125421,0.0042226813,0.0017801234],"domain_scores_gemma":[0.68343306,0.27263707,0.010899252,0.0059666703,0.022233607,0.0048303464],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.020882344,0.0005430647,0.0005120643,0.0030966098,0.0077236993,0.0032719232,0.0019778202,0.0026454711,0.0043702777],"category_scores_gemma":[0.105523,0.00060381956,0.00039893197,0.0023654127,0.004020767,0.003344868,0.0041888044,0.0033336736,0.00071701535],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001463651,0.00029088336,0.010139069,0.0009338329,0.000015389827,0.0026550405,0.9474417,0.00012682051,0.009426288,0.002809386,0.002477184,0.023538016],"study_design_scores_gemma":[0.00004222695,0.00037005052,0.015841687,0.00092673744,0.000031488627,0.0013063436,0.9332808,0.00063039566,0.006296795,0.0018573321,0.039333086,0.000083127634],"about_ca_topic_score_codex":0.0054790406,"about_ca_topic_score_gemma":0.013908016,"teacher_disagreement_score":0.97911763,"about_ca_system_score_codex":0.0058104545,"about_ca_system_score_gemma":0.00716713,"threshold_uncertainty_score":0.11043769},"labels":[],"label_agreement":null},{"id":"W4403976907","doi":"10.1371/journal.pone.0310840","title":"Selecting optimal software code descriptors—The case of Java","year":2024,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Analytical Center for the Government of the Russian Federation","keywords":"Computer science; Java; Overfitting; Software; Class (philosophy); Collinearity; Set (abstract data type); Metric (unit); Data mining; Perspective (graphical); Identification (biology); Software metric; Machine learning; Software development; Focus (optics); Similarity (geometry); Source lines of code; Artificial intelligence; Software quality; Programming language; Mathematics","score_opus":0.05683408453195925,"score_gpt":0.2655602238300151,"score_spread":0.20872613929805583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403976907","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8888668,0.0028287428,0.09921998,0.0010408929,0.000103851664,0.00022086312,0.0030924578,0.0018425114,0.0027840121],"genre_scores_gemma":[0.80643964,0.00052729377,0.18058054,0.00015360146,0.0000439537,0.00021178632,0.010369262,0.00042195976,0.0012519538],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963582,0.0010376349,0.0003225895,0.0010156501,0.0009982684,0.00026756746],"domain_scores_gemma":[0.99465925,0.0025033369,0.00046748912,0.0007739372,0.0012281071,0.00036785036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034315179,0.0012542378,0.001375403,0.0043354193,0.0007147582,0.0015559336,0.0012244193,0.0012149365,0.00050685403],"category_scores_gemma":[0.015245349,0.0003099195,0.0010686503,0.003120458,0.00077853433,0.0016855364,0.0011151169,0.0009518249,0.0004622091],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011207588,0.0011334749,0.09439606,0.0014501914,0.0003575776,0.0012440739,0.0008786474,0.113665394,0.019742545,0.005651178,0.024960976,0.7353991],"study_design_scores_gemma":[0.0003031654,0.0007158423,0.041965958,0.00022956873,0.00023403384,0.0010840424,0.0013760967,0.890119,0.01951015,0.023537757,0.020839822,0.00008452927],"about_ca_topic_score_codex":0.0068448056,"about_ca_topic_score_gemma":0.009308751,"teacher_disagreement_score":0.0068448056,"about_ca_system_score_codex":0.0010430898,"about_ca_system_score_gemma":0.0018153855,"threshold_uncertainty_score":0.018147826},"labels":[],"label_agreement":null},{"id":"W4404029965","doi":"10.1109/icccnt61001.2024.10725742","title":"Innovating Software Complexity Measurement: A New Dimension in OO-Based Metrics","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hatch (Canada)","funders":"","keywords":"Computer science; Dimension (graph theory); Software metric; Software measurement; Software; Software engineering; Programming complexity; Software quality; Software construction; Software development; Programming language; Mathematics","score_opus":0.11350064238837262,"score_gpt":0.3061502340404313,"score_spread":0.19264959165205867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404029965","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0252889,0.0029998736,0.95461035,0.0025263312,0.00035656465,0.00011244882,0.00019635358,0.00054995035,0.013359303],"genre_scores_gemma":[0.390561,0.002387616,0.60205173,0.0006038186,0.0004997549,0.00041432993,0.00047094445,0.0005946892,0.0024161485],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9856021,0.0044289064,0.0013136847,0.001315344,0.0069606877,0.0003793924],"domain_scores_gemma":[0.96264267,0.018841183,0.0046892054,0.0050170515,0.007358802,0.0014510724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008054876,0.0014198087,0.00088131137,0.007781615,0.0011423961,0.00545248,0.0015720101,0.0012801448,0.0015185573],"category_scores_gemma":[0.052188307,0.0005134663,0.00080738665,0.0071305414,0.0051454464,0.014077922,0.0042774673,0.0027764668,0.00042216628],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000082984145,0.00010773052,0.017615372,0.0008497332,0.00014775601,0.0001128722,0.002115016,0.019082071,0.0071446327,0.63245875,0.004602424,0.3156807],"study_design_scores_gemma":[0.000033989854,0.0003864471,0.013933135,0.000616486,0.00011977569,0.00055581104,0.0016476772,0.08199813,0.006946029,0.79006267,0.10344212,0.00025775036],"about_ca_topic_score_codex":0.001888929,"about_ca_topic_score_gemma":0.0015887861,"teacher_disagreement_score":0.008054876,"about_ca_system_score_codex":0.0019273299,"about_ca_system_score_gemma":0.002128349,"threshold_uncertainty_score":0.042598784},"labels":[],"label_agreement":null},{"id":"W4404030792","doi":"10.1109/icccnt61001.2024.10724880","title":"Beyond Traditional Metrics: Advancing Software Complexity Analysis with the ACB Metric","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hatch (Canada)","funders":"","keywords":"Computer science; Metric (unit); Software metric; Software; Software engineering; Software development; Software construction; Programming language; Engineering","score_opus":0.036916693521897186,"score_gpt":0.27417376174763364,"score_spread":0.23725706822573644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404030792","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030029537,0.0018211511,0.9511093,0.001617335,0.00014316903,0.0002218651,0.00036070554,0.00083973637,0.0138571905],"genre_scores_gemma":[0.35983002,0.0015791758,0.63413185,0.00039301973,0.00023323773,0.0005397779,0.00091243454,0.00066388375,0.0017165978],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9778414,0.0067496975,0.0014653951,0.0015692806,0.011975255,0.0003989799],"domain_scores_gemma":[0.8889795,0.058140397,0.013494244,0.0126055535,0.02442854,0.002351689],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011224206,0.0018518972,0.0012451854,0.017765041,0.0013894585,0.006653188,0.001999693,0.0014556845,0.0026324412],"category_scores_gemma":[0.11553938,0.0006564222,0.0010913205,0.01069459,0.0040927045,0.018045632,0.0050427294,0.003325739,0.00081551983],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012628907,0.0002176767,0.03731025,0.0012743085,0.000283553,0.00020887199,0.0031930637,0.034499265,0.0062164124,0.46678957,0.0077070524,0.44217387],"study_design_scores_gemma":[0.000030526968,0.00036898503,0.018079551,0.0006397666,0.00011120816,0.0004433357,0.0016263203,0.19172709,0.004352695,0.72358125,0.058828335,0.00021080546],"about_ca_topic_score_codex":0.0048715486,"about_ca_topic_score_gemma":0.004266115,"teacher_disagreement_score":0.017765041,"about_ca_system_score_codex":0.0026252263,"about_ca_system_score_gemma":0.0037832663,"threshold_uncertainty_score":0.059360027},"labels":[],"label_agreement":null},{"id":"W4404032580","doi":"10.1109/icccnt61001.2024.10725820","title":"Enhancing Software Development with CCC: A Robust Tool for Code Complexity Analysis","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hatch (Canada)","funders":"","keywords":"Computer science; Software engineering; Programming language; Static program analysis; Code (set theory); Software development; Software","score_opus":0.0450906517924193,"score_gpt":0.28452961062054744,"score_spread":0.23943895882812813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404032580","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004623253,0.00016950666,0.97586834,0.00013537686,0.00006443184,0.00019342858,0.0002629693,0.016963782,0.001719002],"genre_scores_gemma":[0.050932717,0.0001917674,0.9441739,0.000100411344,0.00005494306,0.0004757636,0.00071447727,0.0023714893,0.0009845626],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98445016,0.0031107678,0.0010186434,0.0014260407,0.009655822,0.0003386596],"domain_scores_gemma":[0.9622401,0.016865231,0.0047952337,0.005541992,0.009844183,0.0007131834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006683997,0.0015173163,0.0011792914,0.0071893586,0.00093763083,0.0031190503,0.0021420687,0.001127408,0.0039944896],"category_scores_gemma":[0.05766169,0.0010126448,0.0010485147,0.0037385924,0.0012357408,0.0041716527,0.0038590736,0.002368436,0.0016236912],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002471781,0.00027712458,0.0141095845,0.001176852,0.00019127021,0.00038923253,0.0014210588,0.05802639,0.024392257,0.064348705,0.03338791,0.8020324],"study_design_scores_gemma":[0.00013285264,0.00043208152,0.007279947,0.0005746146,0.0001614329,0.0009710863,0.00029508595,0.78737,0.037545796,0.046732724,0.1181236,0.0003807307],"about_ca_topic_score_codex":0.0031625768,"about_ca_topic_score_gemma":0.003393634,"teacher_disagreement_score":0.0071893586,"about_ca_system_score_codex":0.0012093628,"about_ca_system_score_gemma":0.0042760395,"threshold_uncertainty_score":0.035348713},"labels":[],"label_agreement":null},{"id":"W4404060176","doi":"10.1145/3702976","title":"Non-Linear Software Documentation with Interactive Code Examples","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Documentation; Software engineering; Software; Programming language","score_opus":0.06900117343039758,"score_gpt":0.34881224904692565,"score_spread":0.27981107561652807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404060176","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30464303,0.0015260592,0.61228603,0.003494464,0.00027892503,0.00092052243,0.0013146411,0.016134009,0.05940231],"genre_scores_gemma":[0.5309378,0.00054958946,0.4488563,0.0003962431,0.00009142871,0.00064687437,0.0012536722,0.0012658241,0.01600215],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949935,0.0025557089,0.00046447603,0.00041687253,0.0013488508,0.00022058982],"domain_scores_gemma":[0.91436785,0.05893858,0.0049483925,0.013862554,0.006601756,0.0012808546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050445287,0.00060993154,0.00039295116,0.0017758749,0.0009424139,0.002985702,0.0018139352,0.0012623324,0.011231177],"category_scores_gemma":[0.04808603,0.00048131836,0.00037468455,0.001970725,0.0011377914,0.005872825,0.003551854,0.0014008665,0.0027777175],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010488572,0.00094212824,0.0106348675,0.0026833029,0.000043809127,0.0015827736,0.02503895,0.0072576767,0.034830756,0.04756466,0.031365458,0.83700675],"study_design_scores_gemma":[0.00064621045,0.0022588454,0.020002883,0.0033150732,0.00015295201,0.006130728,0.011993823,0.07559513,0.071392246,0.08682447,0.72124547,0.0004420865],"about_ca_topic_score_codex":0.00082590105,"about_ca_topic_score_gemma":0.0027853863,"teacher_disagreement_score":0.011231177,"about_ca_system_score_codex":0.0005792406,"about_ca_system_score_gemma":0.0016163912,"threshold_uncertainty_score":0.037572026},"labels":[],"label_agreement":null},{"id":"W4404165580","doi":"10.1007/978-3-031-73151-8_4","title":"Interpretable SHAP-Driven Machine Learning for Accurate Fault Detection in Software Engineering","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Fault detection and isolation; Software engineering; Artificial intelligence; Machine learning","score_opus":0.013019176892114686,"score_gpt":0.23504802738171549,"score_spread":0.2220288504896008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404165580","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005772181,0.0003189786,0.99118984,0.00017377862,0.00005530877,0.000027506498,0.00015182009,0.001145206,0.0011654252],"genre_scores_gemma":[0.5777202,0.00044493264,0.41626328,0.00025006264,0.00012614549,0.00018140678,0.00095200405,0.00032277292,0.0037391325],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992975,0.00023638627,0.00004536959,0.00014579833,0.00022001918,0.000054892185],"domain_scores_gemma":[0.99653685,0.0024337422,0.00014465558,0.0003934131,0.0004502617,0.000041066152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012207185,0.00076404295,0.0008838788,0.00070551655,0.00029665855,0.0011887194,0.0015616757,0.0010624379,0.003986753],"category_scores_gemma":[0.0074891397,0.00042572914,0.0006327544,0.0006491879,0.0007558906,0.0015722062,0.0012488746,0.002129743,0.0008986378],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020549457,0.00013014184,0.00095115014,0.0002578185,0.000053300646,0.00014573109,0.00012898259,0.6249362,0.0053515006,0.05706082,0.0073428387,0.30343592],"study_design_scores_gemma":[0.0000020476023,0.0000076537435,0.000041535393,0.0000062932586,0.0000022111376,0.000009054093,0.000003637319,0.9802074,0.0005793761,0.01878207,0.0003566348,0.0000021424764],"about_ca_topic_score_codex":0.0017547882,"about_ca_topic_score_gemma":0.0030203594,"teacher_disagreement_score":0.003986753,"about_ca_system_score_codex":0.00082205964,"about_ca_system_score_gemma":0.0007307791,"threshold_uncertainty_score":0.013337016},"labels":[],"label_agreement":null},{"id":"W4404202137","doi":"10.1016/j.jss.2024.112277","title":"Retriever: A view-based approach to reverse engineering software architecture models","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Reverse engineering; Software engineering; Computer science; Architecture; Labrador Retriever; Software; Systems engineering; Engineering; Operating system; Geography; Medicine","score_opus":0.018278120537309806,"score_gpt":0.231958351685927,"score_spread":0.21368023114861717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404202137","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001459824,0.00009428659,0.9818525,0.00012709602,0.000037892205,0.00013463273,0.0004379879,0.013891685,0.0019641651],"genre_scores_gemma":[0.033159032,0.00038660524,0.9547531,0.00020956453,0.000037791306,0.00017370246,0.0030509073,0.005069724,0.0031595055],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99578756,0.0009905171,0.0003236854,0.0005427937,0.0021131057,0.00024227986],"domain_scores_gemma":[0.99178946,0.0027503027,0.00043162418,0.0037754523,0.0011047166,0.00014840068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003988524,0.0019545597,0.001496027,0.004034907,0.0009693104,0.0056254966,0.004667809,0.002422213,0.010620115],"category_scores_gemma":[0.015204713,0.0020537484,0.004721082,0.0024145446,0.0014974909,0.008047681,0.0052121286,0.004401502,0.005090884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003750466,0.0004961226,0.0036847226,0.0014825822,0.00054870767,0.0012483319,0.0018667389,0.06345002,0.022254305,0.30074668,0.03614574,0.5677011],"study_design_scores_gemma":[0.00015393931,0.00021427726,0.0006993478,0.0004971905,0.0003991256,0.0010477905,0.0006564811,0.6106116,0.030331695,0.21780832,0.13738413,0.00019612731],"about_ca_topic_score_codex":0.0064265924,"about_ca_topic_score_gemma":0.01637675,"teacher_disagreement_score":0.010620115,"about_ca_system_score_codex":0.0012640402,"about_ca_system_score_gemma":0.0022027672,"threshold_uncertainty_score":0.035527825},"labels":[],"label_agreement":null},{"id":"W4404260893","doi":"10.48550/arxiv.2410.16469","title":"Evaluating the Performance of a D-Wave Quantum Annealing System for Feature Subset Selection in Software Defect Prediction","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Feature selection; Quantum; Software; Simulated annealing; Quantum annealing; Annealing (glass); Computer science; Feature (linguistics); Selection (genetic algorithm); Artificial intelligence; Pattern recognition (psychology); Materials science; Algorithm; Physics; Quantum computer; Quantum mechanics; Operating system; Composite material","score_opus":0.10341556230274525,"score_gpt":0.24457097089727628,"score_spread":0.14115540859453102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404260893","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67871803,0.00096319465,0.3131765,0.0008532126,0.00012531651,0.00019007543,0.00032051225,0.0017517407,0.0039013498],"genre_scores_gemma":[0.8390459,0.00016655796,0.15893777,0.00020195765,0.000018619614,0.00013896252,0.0005404141,0.00005689701,0.00089285464],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944884,0.00022176995,0.00004003825,0.00011041007,0.00012573524,0.000053273787],"domain_scores_gemma":[0.9971718,0.0019280018,0.00013900636,0.0002069696,0.0004602681,0.00009395376],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022192746,0.0007341862,0.00082390406,0.0007953512,0.0004808382,0.00069944485,0.0011025666,0.0012677967,0.0011279919],"category_scores_gemma":[0.005473408,0.00032775986,0.00071544823,0.00072588125,0.0005100775,0.0010255055,0.0005422134,0.00088492373,0.00022297287],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003649581,0.00024940862,0.004898176,0.00015582294,0.00010770557,0.00006511699,0.00007074695,0.91384846,0.0051097926,0.00304922,0.0014296259,0.07065101],"study_design_scores_gemma":[0.000013780126,0.000049501246,0.0002908928,0.0000022222864,0.0000056467516,0.000005178539,0.000008443423,0.99804556,0.0010909729,0.00038202934,0.000102268816,0.000003433876],"about_ca_topic_score_codex":0.007889559,"about_ca_topic_score_gemma":0.006025869,"teacher_disagreement_score":0.007889559,"about_ca_system_score_codex":0.000717285,"about_ca_system_score_gemma":0.0012972801,"threshold_uncertainty_score":0.015687287},"labels":[],"label_agreement":null},{"id":"W4404371343","doi":"10.1109/icsintesa62455.2024.10748227","title":"A Comprehensive Software Complexity Metric Based on Cyclomatic Complexity","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hatch (Canada)","funders":"","keywords":"Cyclomatic complexity; Computer science; Metric (unit); Software metric; Software; Software engineering; Software quality; Software development; Theoretical computer science; Programming language; Engineering","score_opus":0.0754668234922073,"score_gpt":0.31474984614339196,"score_spread":0.23928302265118467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404371343","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20901518,0.0021085264,0.74867165,0.0007135434,0.00020690337,0.0005560326,0.002254358,0.0012943042,0.03517955],"genre_scores_gemma":[0.7724387,0.0008032262,0.22035566,0.00013358217,0.0001100367,0.00047417518,0.0024390635,0.00018597784,0.0030595835],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9943019,0.0007871585,0.00051732914,0.00044045545,0.0037593597,0.0001936466],"domain_scores_gemma":[0.98621005,0.0039274464,0.0027180377,0.0012573347,0.0051207,0.0007664819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019456067,0.00082169037,0.00046915087,0.008198721,0.0005800245,0.001999905,0.00071949355,0.0005424537,0.0026355153],"category_scores_gemma":[0.016629422,0.00021457141,0.00057778467,0.0042966963,0.0011781316,0.0040291874,0.0018541913,0.00068398035,0.00037894776],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036378653,0.00032138175,0.11829178,0.0015447487,0.00045346044,0.0004290023,0.0015829046,0.081984036,0.039652504,0.16078746,0.014446734,0.5801422],"study_design_scores_gemma":[0.000111467365,0.0024175532,0.27025595,0.0009399505,0.00048438797,0.0035078446,0.0027381226,0.3512314,0.051531408,0.18694142,0.12922329,0.00061713747],"about_ca_topic_score_codex":0.0020234461,"about_ca_topic_score_gemma":0.0024613475,"teacher_disagreement_score":0.008198721,"about_ca_system_score_codex":0.0012951943,"about_ca_system_score_gemma":0.0013967089,"threshold_uncertainty_score":0.010289431},"labels":[],"label_agreement":null},{"id":"W4404635154","doi":"10.1145/3705309","title":"Detecting Refactoring Commits in Machine Learning Python Projects: A Machine Learning-Based Approach","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Code refactoring; Computer science; Maintainability; Python (programming language); Software engineering; Java; Programming language; Software; Software development; Artificial intelligence; Machine learning","score_opus":0.09068230133371412,"score_gpt":0.3197279064261304,"score_spread":0.22904560509241628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404635154","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4966501,0.0038665126,0.439967,0.0014069127,0.000501274,0.0015082407,0.01315624,0.036248405,0.006695309],"genre_scores_gemma":[0.6836524,0.0006322105,0.28518766,0.0003896836,0.0002009567,0.00065166125,0.023776462,0.0005354983,0.004973387],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99225205,0.000958468,0.0009818639,0.0026146893,0.002451739,0.0007411957],"domain_scores_gemma":[0.9791443,0.0071102646,0.004693426,0.0024046514,0.0056228535,0.001024528],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004857549,0.001568887,0.001418849,0.01087656,0.0010773907,0.0021334868,0.002880781,0.0018624122,0.0011272583],"category_scores_gemma":[0.018672936,0.0005176634,0.0014218369,0.004731525,0.0006450117,0.0023773387,0.0018900326,0.0019053183,0.0018734427],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005643325,0.0012942326,0.29552376,0.0007988716,0.00034831036,0.00154814,0.00086148805,0.026482167,0.010463891,0.0018818207,0.021855572,0.63837737],"study_design_scores_gemma":[0.00007068383,0.00038134726,0.08840367,0.00023684306,0.00026738815,0.0012011067,0.0007741007,0.87018025,0.018058028,0.00625794,0.014020757,0.00014791946],"about_ca_topic_score_codex":0.008461891,"about_ca_topic_score_gemma":0.012582295,"teacher_disagreement_score":0.01087656,"about_ca_system_score_codex":0.0010998694,"about_ca_system_score_gemma":0.0024633294,"threshold_uncertainty_score":0.025689542},"labels":[],"label_agreement":null},{"id":"W4404783983","doi":"10.18653/v1/2024.emnlp-industry.68","title":"A new approach for fine-tuning sentence transformers for intent classification and out-of-scope detection tasks","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Mitacs","keywords":"Computer science; Scope (computer science); Sentence; Transformer; Artificial intelligence; Natural language processing; Engineering; Electrical engineering; Programming language; Voltage","score_opus":0.06658125175465925,"score_gpt":0.3089258718649948,"score_spread":0.24234462011033558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404783983","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048034064,0.00057173683,0.9346167,0.00035786311,0.00018659959,0.0002725296,0.00038077423,0.013579702,0.0020000073],"genre_scores_gemma":[0.50135,0.00029731856,0.4890442,0.0006433529,0.00022335374,0.00029363102,0.0019068482,0.0010279431,0.0052134204],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99784887,0.0005502533,0.0002255964,0.0005798209,0.0006596601,0.00013574275],"domain_scores_gemma":[0.9954743,0.0014425667,0.00039986687,0.0011321396,0.0013342066,0.00021682894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002563236,0.0014424688,0.0008630226,0.0015364607,0.0004564982,0.0014119665,0.0014991324,0.00096213917,0.003054308],"category_scores_gemma":[0.010605436,0.0005108962,0.00081386784,0.00072155317,0.0007606445,0.0045747696,0.0020547416,0.0024283251,0.0024951536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006557112,0.00067232514,0.0055178395,0.00024965272,0.0001283002,0.00016082909,0.00042530466,0.023562815,0.09298137,0.005028963,0.010436096,0.8601808],"study_design_scores_gemma":[0.00008049019,0.0006950927,0.0024957643,0.0000382489,0.00009240998,0.00054521963,0.00018162977,0.9187397,0.060744334,0.009180708,0.0071483203,0.000057971865],"about_ca_topic_score_codex":0.0017810091,"about_ca_topic_score_gemma":0.0031612753,"teacher_disagreement_score":0.003054308,"about_ca_system_score_codex":0.00063661824,"about_ca_system_score_gemma":0.0013266501,"threshold_uncertainty_score":0.013555884},"labels":[],"label_agreement":null},{"id":"W4404800967","doi":"10.1016/j.jss.2024.112283","title":"RPerf: Mining user reviews using topic modeling to assist performance testing: An industrial experience report","year":2024,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Engineering; Software engineering; Data science","score_opus":0.19399540972428775,"score_gpt":0.34820829008042375,"score_spread":0.154212880356136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404800967","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5356715,0.013918531,0.20969872,0.002991868,0.0012843761,0.0034026166,0.08703903,0.13322042,0.012772933],"genre_scores_gemma":[0.49839637,0.0031797907,0.33950737,0.00060690497,0.00096876733,0.0019502083,0.12847559,0.0033469042,0.023567993],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9929315,0.0023954646,0.0005724883,0.0010662746,0.0027758398,0.00025840686],"domain_scores_gemma":[0.9662524,0.020393861,0.0020180328,0.0029031888,0.007282561,0.0011499774],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007943445,0.0019312119,0.0016184287,0.007813564,0.00087369897,0.0023865066,0.0020837253,0.0013597116,0.0031032884],"category_scores_gemma":[0.02642595,0.00049742585,0.001161577,0.0039041408,0.00027402115,0.0031916136,0.001323381,0.0011461818,0.006611558],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001050572,0.0025135127,0.03749338,0.0021806466,0.0008810499,0.00043131856,0.0016555985,0.0042553586,0.022747885,0.00058586564,0.16296951,0.7632354],"study_design_scores_gemma":[0.0008586553,0.0052620936,0.19045843,0.00056345505,0.0019194321,0.002861753,0.002650632,0.5410689,0.07088799,0.0027658807,0.18002823,0.00067450956],"about_ca_topic_score_codex":0.005712152,"about_ca_topic_score_gemma":0.0079829935,"teacher_disagreement_score":0.007943445,"about_ca_system_score_codex":0.00049671205,"about_ca_system_score_gemma":0.0012628555,"threshold_uncertainty_score":0.042009473},"labels":[],"label_agreement":null},{"id":"W4404837934","doi":"10.1016/j.procs.2024.09.568","title":"Automating Software Documentation: Employing LLMs for Precise Use Case Description","year":2024,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"Centre National pour la Recherche Scientifique et Technique","keywords":"Computer science; Documentation; Software; Software engineering; Software documentation; Data science; Programming language; Software development; Software development process","score_opus":0.04452629760129762,"score_gpt":0.31453133743875383,"score_spread":0.2700050398374562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404837934","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0072322637,0.00008798126,0.9873104,0.00015230544,0.000010105997,0.0003433792,0.00013324081,0.0026983942,0.0020318513],"genre_scores_gemma":[0.052628815,0.000101596306,0.9449503,0.00003714295,0.000006942165,0.00036855412,0.00042591122,0.00032765474,0.0011531665],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9910978,0.005132962,0.0007698478,0.00081075425,0.0019949183,0.0001937106],"domain_scores_gemma":[0.9743882,0.015488009,0.0019391931,0.0058893557,0.0020663086,0.00022900205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072208648,0.0009073746,0.0005132755,0.0052161505,0.0009077958,0.0033293986,0.0016883615,0.0012436173,0.003203977],"category_scores_gemma":[0.030026976,0.00086753746,0.0010369509,0.002169327,0.0011948398,0.003585062,0.0030906335,0.0015147496,0.001996235],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013374975,0.00029990455,0.0048610237,0.0009872003,0.00006076523,0.0014663336,0.012309444,0.021238677,0.04568926,0.053998888,0.0050808717,0.85387385],"study_design_scores_gemma":[0.0001159936,0.00031997857,0.004609942,0.0016526367,0.0001303732,0.0038513117,0.0042771236,0.6317528,0.11097854,0.08184901,0.16022712,0.00023520837],"about_ca_topic_score_codex":0.0018713449,"about_ca_topic_score_gemma":0.0038516133,"teacher_disagreement_score":0.0072208648,"about_ca_system_score_codex":0.0012705178,"about_ca_system_score_gemma":0.0028207682,"threshold_uncertainty_score":0.03818804},"labels":[],"label_agreement":null},{"id":"W4404952914","doi":"10.1109/issre62328.2024.00030","title":"Assessing the Performance of AI-Generated Code: A Case Study on GitHub Copilot","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Natural Science Foundation of China","keywords":"Computer science; Code (set theory); Computer security; Programming language","score_opus":0.06393546958244994,"score_gpt":0.3681672486282733,"score_spread":0.3042317790458234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404952914","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9734344,0.00047125094,0.0098058395,0.0007333914,0.00006420976,0.00043309314,0.0034799664,0.0074258577,0.004152032],"genre_scores_gemma":[0.9131528,0.00040491138,0.060307547,0.0005953745,0.000037271122,0.0007368449,0.01774064,0.0036770361,0.003347513],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.995285,0.0015601082,0.00026124864,0.00075086893,0.0018858758,0.000256962],"domain_scores_gemma":[0.96657604,0.02321876,0.0017007291,0.003951603,0.0038555132,0.0006973378],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004273068,0.00095422514,0.00044066276,0.0020791919,0.00087787467,0.0011642954,0.002692188,0.0012025165,0.0010253398],"category_scores_gemma":[0.027321491,0.00041798525,0.0006385326,0.002709116,0.0022292922,0.0013403,0.0015242253,0.0018222764,0.0008817649],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0054502185,0.0052780258,0.12898207,0.005672443,0.00060857995,0.014048368,0.021489246,0.19397463,0.05073551,0.011143199,0.11469232,0.44792536],"study_design_scores_gemma":[0.0012634356,0.004381435,0.22976623,0.00087474234,0.00029465006,0.0040249396,0.006433326,0.5483323,0.083205774,0.010890966,0.11006013,0.00047203348],"about_ca_topic_score_codex":0.012488953,"about_ca_topic_score_gemma":0.019419389,"teacher_disagreement_score":0.012488953,"about_ca_system_score_codex":0.0021789218,"about_ca_system_score_gemma":0.0018306988,"threshold_uncertainty_score":0.024832547},"labels":[],"label_agreement":null},{"id":"W4405098729","doi":"10.22215/etd/2024-16309","title":"Beyond Verbal Self-Explanations: Student Annotations of a Code-Tracing Example Produced by ChatGPT","year":2024,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Flowchart; Tracing; Computer science; Process tracing; Code (set theory); Constructive; Annotation; Expression (computer science); Quality (philosophy); Variety (cybernetics); Domain (mathematical analysis); Process (computing); Mathematics education; Programming language; Natural language processing; Artificial intelligence; Psychology; Mathematics","score_opus":0.013385135261252052,"score_gpt":0.3064894869946401,"score_spread":0.29310435173338806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405098729","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8505205,0.00021760054,0.133465,0.0014333696,0.00011803531,0.00035205262,0.00040653386,0.002268951,0.011217976],"genre_scores_gemma":[0.8900295,0.00017911145,0.10280382,0.00021418155,0.000027674669,0.00033761412,0.0004142168,0.00050260697,0.005491271],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.98309696,0.012106248,0.0006911355,0.0011747197,0.0025236434,0.00040735793],"domain_scores_gemma":[0.81294894,0.15196852,0.009390515,0.009236725,0.014886576,0.0015687156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013111895,0.0008257153,0.0005499565,0.0019517053,0.0013454638,0.003096701,0.0013966836,0.0015714152,0.0034299935],"category_scores_gemma":[0.11374516,0.0003455662,0.00045437377,0.0013677458,0.0016958598,0.0034488824,0.0034306129,0.0017946443,0.0011376077],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010999562,0.00061723054,0.06978074,0.0016265906,0.000086577056,0.0018434832,0.5084612,0.005484608,0.051091332,0.004606869,0.0068887156,0.3484127],"study_design_scores_gemma":[0.00025753985,0.003503942,0.13209423,0.0041439654,0.00044505834,0.004165375,0.33499035,0.13468432,0.19114485,0.023767943,0.1696933,0.0011091318],"about_ca_topic_score_codex":0.0012417968,"about_ca_topic_score_gemma":0.0027571556,"teacher_disagreement_score":0.013111895,"about_ca_system_score_codex":0.0010743196,"about_ca_system_score_gemma":0.0012671497,"threshold_uncertainty_score":0.06934315},"labels":[],"label_agreement":null},{"id":"W4405180284","doi":"10.1109/ict4da62874.2024.10777285","title":"Software Complexity Analysis with the Advanced CB Metric","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hatch (Canada)","funders":"","keywords":"Computer science; Metric (unit); Software; Software metric; Software engineering; Programming language; Software construction; Software development; Engineering","score_opus":0.023458573983166822,"score_gpt":0.2832560903920291,"score_spread":0.2597975164088623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405180284","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09129393,0.00080081215,0.89010674,0.00037220967,0.0000537029,0.00038739797,0.0006456202,0.0008842428,0.015455295],"genre_scores_gemma":[0.5934818,0.00042922638,0.40158743,0.000093282324,0.00006422319,0.0005784339,0.001047395,0.00025866835,0.0024594956],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99211544,0.0015564203,0.0005236835,0.000658531,0.004843648,0.00030232695],"domain_scores_gemma":[0.9702631,0.014766686,0.003955821,0.0030395763,0.0072027016,0.0007721367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003449592,0.0009698716,0.00053676456,0.010501472,0.00069820805,0.002364948,0.0009845729,0.0006420558,0.0027572475],"category_scores_gemma":[0.03609567,0.00029684132,0.00094982825,0.005446448,0.0016033158,0.004928054,0.0023043433,0.0012124815,0.00043880247],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003478852,0.00023450762,0.06723927,0.0009328376,0.00027315068,0.00044657412,0.0025554814,0.14874995,0.019676087,0.31862092,0.008405061,0.4325183],"study_design_scores_gemma":[0.000058047815,0.00065766537,0.05112285,0.0003043345,0.00015381178,0.00086720934,0.0013562952,0.6559771,0.011469445,0.24211158,0.03571502,0.00020657606],"about_ca_topic_score_codex":0.0068220748,"about_ca_topic_score_gemma":0.004184061,"teacher_disagreement_score":0.010501472,"about_ca_system_score_codex":0.0018771985,"about_ca_system_score_gemma":0.0020139348,"threshold_uncertainty_score":0.018243372},"labels":[],"label_agreement":null},{"id":"W4405259287","doi":"10.1007/s10664-024-10597-8","title":"Harnessing pre-trained generalist agents for software engineering tasks","year":2024,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Reinforcement learning; Software engineering; Generalizability theory; Scheduling (production processes); Machine learning; Software; Domain (mathematical analysis); Human–computer interaction; Engineering; Operations management","score_opus":0.03556022016790315,"score_gpt":0.3148634202315637,"score_spread":0.27930320006366055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405259287","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.62802964,0.0003782523,0.32762682,0.0007724147,0.00019025734,0.00051031227,0.00018534489,0.0057170033,0.036589894],"genre_scores_gemma":[0.8741548,0.00019760775,0.1099425,0.0004115148,0.00004997045,0.00018795808,0.0003745756,0.00026055184,0.014420534],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993656,0.00019909507,0.000030487645,0.00019339516,0.00013585186,0.00007560098],"domain_scores_gemma":[0.99622923,0.0014999069,0.0002870033,0.0009911491,0.0006378599,0.00035477342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013008745,0.0006634757,0.00044571763,0.00045406283,0.0004292776,0.0011677525,0.0012311325,0.0010687689,0.0060692984],"category_scores_gemma":[0.007751285,0.00046216237,0.00029245156,0.00038313726,0.00048056876,0.0013012084,0.0020978525,0.0014273361,0.0028948314],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011413154,0.002910708,0.03333971,0.00054472656,0.00024353249,0.000576374,0.0032179414,0.076391645,0.1258015,0.0066806544,0.009739084,0.7394127],"study_design_scores_gemma":[0.00031067137,0.0024782894,0.029582808,0.00016897355,0.00031075155,0.00062811654,0.0020905721,0.8290638,0.06173095,0.020788083,0.05270581,0.00014122254],"about_ca_topic_score_codex":0.0021513454,"about_ca_topic_score_gemma":0.0059024733,"teacher_disagreement_score":0.0060692984,"about_ca_system_score_codex":0.00047066194,"about_ca_system_score_gemma":0.0011421628,"threshold_uncertainty_score":0.020303845},"labels":[],"label_agreement":null},{"id":"W4405396136","doi":"10.1145/3708519","title":"Automatic Programming: Large Language Models and Beyond","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Bộ Giáo dục và Ðào tạo; Ministry of Education, India","keywords":"Computer science; Programmer; Coding (social sciences); Software engineering; Popularity; Software deployment; Programming language; Key (lock); Computer security","score_opus":0.05968806315916537,"score_gpt":0.33576718525188126,"score_spread":0.2760791220927159,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405396136","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009965851,0.0014914011,0.97095245,0.0037583855,0.00009565885,0.00010418605,0.00041044346,0.0022298843,0.010991781],"genre_scores_gemma":[0.2633992,0.0031765634,0.71714664,0.0014653134,0.00047801764,0.00078109285,0.0014204689,0.0024556576,0.009677043],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99477625,0.002345012,0.0002466421,0.0007361267,0.0016309092,0.00026500938],"domain_scores_gemma":[0.9757829,0.016439559,0.0011723186,0.004596164,0.001584777,0.0004243088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004962765,0.00083473226,0.0009310224,0.0015368704,0.0012856335,0.005560269,0.0031098612,0.0018232376,0.005819689],"category_scores_gemma":[0.020628989,0.0011929276,0.0020523826,0.0015249613,0.004078828,0.012480675,0.0035846105,0.0049659875,0.0017139299],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000047998696,0.00005925855,0.001097725,0.00022632678,0.00004769448,0.00016624862,0.0009206434,0.0326223,0.0011590582,0.9262631,0.0048767123,0.032512926],"study_design_scores_gemma":[0.00002006936,0.000018408566,0.00022143484,0.00011852688,0.000019063198,0.0001166784,0.00010356177,0.1339834,0.0006784028,0.834369,0.030325301,0.00002618871],"about_ca_topic_score_codex":0.004570557,"about_ca_topic_score_gemma":0.0042747064,"teacher_disagreement_score":0.005819689,"about_ca_system_score_codex":0.0020745185,"about_ca_system_score_gemma":0.0027021372,"threshold_uncertainty_score":0.026245952},"labels":[],"label_agreement":null},{"id":"W4405522649","doi":"10.1109/conisoft63288.2024.00021","title":"Next-Generation Metric for Software Complexity Analysis","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hatch (Canada)","funders":"","keywords":"Computer science; Metric (unit); Software; Software metric; Software engineering; Theoretical computer science; Software development; Programming language; Software construction; Engineering","score_opus":0.12728627278484023,"score_gpt":0.33231210303461045,"score_spread":0.20502583024977022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405522649","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047960497,0.0016073552,0.9141066,0.000662935,0.00040211304,0.0006322541,0.0024769946,0.0026758343,0.029475411],"genre_scores_gemma":[0.38276497,0.0006384669,0.60354406,0.00024993916,0.00015219382,0.0016095387,0.004817,0.00052954955,0.005694354],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9905005,0.0021528418,0.00068699196,0.00091367465,0.005371831,0.00037428213],"domain_scores_gemma":[0.97950256,0.0076482664,0.0025198718,0.0026678822,0.0068167145,0.0008446597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049434053,0.0013525094,0.00068229024,0.00999846,0.0009893181,0.002638061,0.0012737414,0.0009797614,0.0058180895],"category_scores_gemma":[0.034661055,0.00024361022,0.0011243075,0.0054229656,0.0011494933,0.0048573622,0.00228496,0.0014012287,0.0015348727],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034807937,0.0003329233,0.045800388,0.0008695722,0.00030782475,0.00031634766,0.001370729,0.06532924,0.01194284,0.21233115,0.02205132,0.6389995],"study_design_scores_gemma":[0.000080888574,0.0010810782,0.039075624,0.00045609724,0.0001465751,0.0014765288,0.0011130865,0.54761255,0.019926935,0.21052164,0.1781147,0.00039435277],"about_ca_topic_score_codex":0.0043083136,"about_ca_topic_score_gemma":0.0035627848,"teacher_disagreement_score":0.00999846,"about_ca_system_score_codex":0.0028081161,"about_ca_system_score_gemma":0.0019603102,"threshold_uncertainty_score":0.026143491},"labels":[],"label_agreement":null},{"id":"W4405601553","doi":"10.1109/scam63643.2024.00024","title":"On the Prevalence, Evolution, and Impact of Code Smells in Simulation Modelling Software","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Software evolution; Code (set theory); Software; Software engineering; Programming language; Software construction; Software development","score_opus":0.028226155708876457,"score_gpt":0.3051652247556653,"score_spread":0.2769390690467889,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405601553","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99851054,0.000120476085,0.00081722246,0.000062752355,0.0000031166196,0.000011189372,0.00006317508,0.000049653132,0.00036197962],"genre_scores_gemma":[0.9983683,0.000077124,0.0011890156,0.000019972122,0.0000044063763,0.000009932602,0.00017623966,0.000022345946,0.00013267614],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98955476,0.0034770772,0.0010697126,0.0016127174,0.003638511,0.0006472521],"domain_scores_gemma":[0.6831837,0.21915694,0.06684391,0.011040645,0.01575414,0.004020631],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008780134,0.00038794088,0.00034175004,0.0046439064,0.0006902533,0.0016553577,0.00061501743,0.00087271,0.0008076394],"category_scores_gemma":[0.1067807,0.00042689688,0.0005109975,0.0022348429,0.0014539819,0.002970791,0.0016457324,0.00129104,0.00019413963],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028537863,0.0002198176,0.9598534,0.00013864673,0.00008532601,0.00031392337,0.0027556017,0.002016845,0.004233518,0.00027708107,0.00022344814,0.029597074],"study_design_scores_gemma":[0.00001076637,0.0004989349,0.979973,0.00006462533,0.00006356768,0.00051888806,0.0014843534,0.013806403,0.0027422113,0.0003243957,0.00046819082,0.00004464652],"about_ca_topic_score_codex":0.0034975754,"about_ca_topic_score_gemma":0.005805773,"teacher_disagreement_score":0.008780134,"about_ca_system_score_codex":0.0009627596,"about_ca_system_score_gemma":0.0006723823,"threshold_uncertainty_score":0.046434343},"labels":[],"label_agreement":null},{"id":"W4405602376","doi":"10.1109/scam63643.2024.00013","title":"AUTOGENICS: Automated Generation of Context-Aware Inline Comments for Code Snippets on Programming Q&amp;A Sites Using LLM","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Context (archaeology); Code generation; Code (set theory); Programming language; Information retrieval; Operating system; Biology","score_opus":0.14257578200308108,"score_gpt":0.38368635805457546,"score_spread":0.24111057605149439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405602376","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13401611,0.00069602847,0.41697225,0.0010195103,0.00049362966,0.0020561893,0.012176908,0.4253839,0.0071853832],"genre_scores_gemma":[0.2390036,0.0004070731,0.7023533,0.0006062345,0.0001554092,0.0020060025,0.01865771,0.02356005,0.013250548],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99645895,0.0014910328,0.00029179768,0.00074384466,0.00085132365,0.0001630589],"domain_scores_gemma":[0.9759134,0.013843889,0.0028888248,0.0026686843,0.003915724,0.00076951686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034725678,0.002488671,0.00079283485,0.003079346,0.0006442392,0.0011660813,0.0015126589,0.0010701678,0.010482288],"category_scores_gemma":[0.023712184,0.0006670401,0.0008543342,0.0009227915,0.00053184264,0.0019476662,0.0023180963,0.0009608993,0.00764515],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016173564,0.0009842805,0.017974447,0.003918744,0.00017248082,0.0018142611,0.010141849,0.0053986534,0.0807338,0.0029855478,0.12758557,0.746673],"study_design_scores_gemma":[0.001155924,0.002450149,0.038729023,0.00150991,0.0003141793,0.002474806,0.007152466,0.35916477,0.22369388,0.0106754815,0.35203058,0.00064877956],"about_ca_topic_score_codex":0.0013384899,"about_ca_topic_score_gemma":0.002877876,"teacher_disagreement_score":0.010482288,"about_ca_system_score_codex":0.0005645423,"about_ca_system_score_gemma":0.0014756739,"threshold_uncertainty_score":0.035066724},"labels":[],"label_agreement":null},{"id":"W4405713656","doi":"10.1007/978-3-031-71769-7_5","title":"Teaching Software Metrology: The Science of Measurement for Software Engineering","year":2024,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Metrology; Software engineering; Software; Computer science; Systems engineering; Engineering; Operating system; Physics","score_opus":0.03775517379103521,"score_gpt":0.2640138087159687,"score_spread":0.22625863492493348,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405713656","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009017901,0.026013901,0.05934415,0.004648244,0.003125531,0.000062690524,0.00030805197,0.0012717079,0.904324],"genre_scores_gemma":[0.006313716,0.011475352,0.021181682,0.0011756526,0.00074678083,0.00005331922,0.00027218318,0.0006145158,0.95816684],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996159,0.00004407019,0.0000104453,0.00004873113,0.0002595916,0.000021254406],"domain_scores_gemma":[0.9995813,0.00017294846,0.00001691755,0.000039956958,0.00014043458,0.000048369784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00035176272,0.0009269874,0.0006297433,0.0015205808,0.0010984391,0.003933847,0.0010695787,0.0012891446,0.06809928],"category_scores_gemma":[0.0012074509,0.00047301516,0.00041224615,0.0021243475,0.001134665,0.0037241827,0.0012663413,0.002962833,0.035757408],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000120763425,0.00005050882,0.00010120138,0.00028032213,0.000004683947,0.00004540095,0.00041470016,0.00068945956,0.0015825126,0.2254768,0.43558905,0.3357533],"study_design_scores_gemma":[0.000001903574,0.000009805288,0.00013565103,0.0001299491,0.0000024383817,0.000102288,0.00007421468,0.0006947682,0.0005127611,0.055676192,0.9426531,0.0000069219677],"about_ca_topic_score_codex":0.002188262,"about_ca_topic_score_gemma":0.0051074764,"teacher_disagreement_score":0.06809928,"about_ca_system_score_codex":0.0014041729,"about_ca_system_score_gemma":0.0015653804,"threshold_uncertainty_score":0.2278148},"labels":[],"label_agreement":null},{"id":"W4405713677","doi":"10.1007/978-3-031-71769-7_12","title":"Teaching Mining Software Repositories","year":2024,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; University of British Columbia; University of Saskatchewan","funders":"","keywords":"Computer science; Software engineering; World Wide Web; Data science","score_opus":0.018720423072953542,"score_gpt":0.2598501445017863,"score_spread":0.24112972142883274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405713677","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007009875,0.0168127,0.3658112,0.009659597,0.0019933209,0.00030111638,0.0023349721,0.00839101,0.5876861],"genre_scores_gemma":[0.019688766,0.013177754,0.13772139,0.001409994,0.0010973768,0.00012658119,0.0044529485,0.00174479,0.8205804],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991642,0.00008665032,0.000030816107,0.000189027,0.0004894966,0.00003972316],"domain_scores_gemma":[0.9975994,0.001227648,0.00010391986,0.00034496747,0.00052275445,0.00020129855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010930424,0.0010043498,0.00081607833,0.004995406,0.0009949844,0.0042013223,0.0017451838,0.001073856,0.06029113],"category_scores_gemma":[0.004522947,0.0009203513,0.0007993775,0.006360767,0.00085762935,0.007192606,0.0018683368,0.0022477412,0.044182595],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000010688735,0.00007499962,0.0006866169,0.00026215563,0.000013130765,0.000045592413,0.00029165973,0.0008395843,0.0014075283,0.025827562,0.2578802,0.71266025],"study_design_scores_gemma":[0.0000084706435,0.00003808445,0.0026870149,0.0004273222,0.000023347258,0.0005538827,0.00040132485,0.006089678,0.003979423,0.095854625,0.889911,0.0000258615],"about_ca_topic_score_codex":0.0023459073,"about_ca_topic_score_gemma":0.008304654,"teacher_disagreement_score":0.06029113,"about_ca_system_score_codex":0.0014130541,"about_ca_system_score_gemma":0.002059897,"threshold_uncertainty_score":0.20169389},"labels":[],"label_agreement":null},{"id":"W4405746722","doi":"10.1016/j.neunet.2024.107067","title":"Promises and perils of using Transformer-based models for SE research","year":2024,"lang":"en","type":"review","venue":"Neural Networks","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Fundamental Research Funds for the Central Universities; Sun Yat-sen University","keywords":"Transformer; Computer science; Artificial intelligence; Machine learning; Engineering; Electrical engineering; Voltage","score_opus":0.3153174685816854,"score_gpt":0.45324801306612184,"score_spread":0.13793054448443642,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405746722","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020875385,0.46251613,0.473913,0.019846136,0.0011366797,0.00017171053,0.0011081662,0.0020825926,0.018350234],"genre_scores_gemma":[0.31051672,0.44376966,0.22672093,0.0037494567,0.0012396093,0.00046598975,0.0031224221,0.0005154828,0.009899816],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99830854,0.0007542619,0.00010934021,0.00032358983,0.00043342664,0.00007079053],"domain_scores_gemma":[0.9886615,0.008325265,0.0002468571,0.000998574,0.0016116812,0.00015600977],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0057215514,0.0013027007,0.00096368184,0.0020217462,0.0002420396,0.0022946661,0.0024415057,0.0015440071,0.0019388131],"category_scores_gemma":[0.015290459,0.00055532507,0.000979104,0.0021300064,0.00080346206,0.0066788727,0.0010755023,0.0031231942,0.0018425084],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000100160156,0.000107773805,0.0032241985,0.0031127487,0.0003835575,0.00007821368,0.00019451433,0.06841533,0.0014277097,0.05254054,0.013184225,0.8572311],"study_design_scores_gemma":[0.00005390933,0.00054531434,0.0034256023,0.0038205518,0.00056607294,0.0004793826,0.00031964373,0.57728493,0.005818314,0.21337767,0.19416106,0.00014757912],"about_ca_topic_score_codex":0.0047532413,"about_ca_topic_score_gemma":0.0077007427,"teacher_disagreement_score":0.99427843,"about_ca_system_score_codex":0.0015342779,"about_ca_system_score_gemma":0.0023691806,"threshold_uncertainty_score":0.030258834},"labels":[],"label_agreement":null},{"id":"W4405836525","doi":"10.1007/s00521-024-10906-8","title":"Correction: Software effort estimation using convolutional neural network and fuzzy clustering","year":2024,"lang":"en","type":"article","venue":"Neural Computing and Applications","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Computational Science and Engineering; Convolutional neural network; Cluster analysis; Artificial intelligence; Estimation; Software; Fuzzy logic; Machine learning; Data mining; Programming language","score_opus":0.021942528856182295,"score_gpt":0.2932689784816275,"score_spread":0.27132644962544517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405836525","genre_codex":"editorial","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012172708,0.0033221557,0.06674712,0.046669938,0.7906609,0.00022292859,0.050192796,0.019109173,0.010902299],"genre_scores_gemma":[0.35712495,0.003686303,0.19228451,0.0319464,0.05885286,0.0009601749,0.03566457,0.016970074,0.3025102],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99155504,0.0010374172,0.0018249116,0.0020444416,0.0028156803,0.0007224617],"domain_scores_gemma":[0.89432114,0.025783036,0.006876571,0.013025635,0.057887986,0.002105702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057065864,0.0028140456,0.002576375,0.0073532886,0.0026087775,0.0035582697,0.004517255,0.0053076795,0.13364545],"category_scores_gemma":[0.13912337,0.0012717636,0.0020001165,0.00789694,0.0013207556,0.0033733502,0.0024837519,0.005396394,0.040875122],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045218383,0.000044610133,0.0037597562,0.0011013603,0.00033182296,0.00092434045,0.00024001655,0.0014033815,0.0010833875,0.0048124255,0.93196213,0.053884584],"study_design_scores_gemma":[0.0007022994,0.0001942521,0.035050407,0.0018093436,0.0007884245,0.0034532722,0.0007829233,0.04424209,0.014197443,0.0401671,0.8581271,0.0004853578],"about_ca_topic_score_codex":0.017381862,"about_ca_topic_score_gemma":0.026526826,"teacher_disagreement_score":0.13364545,"about_ca_system_score_codex":0.002153123,"about_ca_system_score_gemma":0.005286068,"threshold_uncertainty_score":0.44708854},"labels":[],"label_agreement":null},{"id":"W4405868299","doi":"10.1002/spe.3401","title":"Android Source Code Smells: A Systematic Literature Review","year":2024,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Code smell; Code refactoring; Android (operating system); Computer science; Software engineering; Software quality; Code review; Software; Data science; World Wide Web; Software development; Programming language; Operating system","score_opus":0.014181815966298993,"score_gpt":0.3108355000391105,"score_spread":0.2966536840728115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405868299","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004833652,0.9906668,0.000747829,0.00074064167,0.0001574895,0.00083307934,0.0010218814,0.000033023116,0.00096565555],"genre_scores_gemma":[0.03382388,0.96005857,0.002648619,0.00058054767,0.00008266486,0.0015584007,0.0009492703,0.000027402946,0.00027058637],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9821191,0.005519872,0.0064387065,0.0013687015,0.004160201,0.0003933859],"domain_scores_gemma":[0.8474994,0.110371694,0.01878622,0.0025315494,0.019911723,0.0008994949],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016820615,0.00154505,0.0042091976,0.03572023,0.00092916266,0.002759006,0.0022236493,0.0016742861,0.0031520952],"category_scores_gemma":[0.08716566,0.0012524418,0.0041014575,0.021911558,0.0015933007,0.0035411608,0.0025771398,0.001106467,0.0005838917],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014164239,0.000028535525,0.0024059133,0.8385756,0.00196264,0.00044074433,0.0015023936,0.0001359802,0.00042533746,0.0004459259,0.0039889654,0.14994636],"study_design_scores_gemma":[0.00007925973,0.00020831675,0.010098664,0.9399859,0.011155392,0.0008376052,0.0019055506,0.00012436825,0.00047622327,0.00037713413,0.034693576,0.000057889632],"about_ca_topic_score_codex":0.010235242,"about_ca_topic_score_gemma":0.030151887,"teacher_disagreement_score":0.03572023,"about_ca_system_score_codex":0.0047953036,"about_ca_system_score_gemma":0.021713652,"threshold_uncertainty_score":0.08895695},"labels":[],"label_agreement":null},{"id":"W4406012448","doi":"10.1109/access.2024.3525069","title":"Leveraging an Enhanced CodeBERT-Based Model for Multiclass Software Defect Prediction via Defect Classification","year":2025,"lang":"en","type":"article","venue":"IEEE Access","topic":"Software Engineering Research","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Regina","funders":"Natural Sciences and Engineering Research Council of Canada; Università degli Studi di Firenze","keywords":"Computer science; Machine learning; Software bug; Artificial intelligence; Software reliability testing; Software development; Software; Software quality; Software construction; Software development process; Context (archaeology); Software engineering; Data mining; Programming language","score_opus":0.05998664035314283,"score_gpt":0.34532641882646503,"score_spread":0.2853397784733222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406012448","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13437405,0.00046515142,0.856349,0.0006681312,0.00007613372,0.00016545641,0.0005549111,0.005552502,0.001794653],"genre_scores_gemma":[0.8252781,0.0002167494,0.16727422,0.00034222545,0.00008469592,0.00028043453,0.0016111142,0.00024041029,0.0046720547],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992944,0.00014772385,0.000043140986,0.00025807245,0.00017916948,0.00007751002],"domain_scores_gemma":[0.996806,0.001746264,0.00035141964,0.00026909471,0.00072132045,0.000105995336],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016849844,0.0010988531,0.0010000545,0.002276528,0.000513208,0.0012461339,0.002258607,0.0015864461,0.0011610278],"category_scores_gemma":[0.0049661244,0.00045663942,0.0010037003,0.0011481423,0.0005816962,0.0013309821,0.0009214069,0.001748216,0.00080457603],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024420486,0.0003704351,0.017514437,0.00006410945,0.000118152166,0.00016648084,0.000117677846,0.78379446,0.003404623,0.0024139513,0.0029988638,0.18879256],"study_design_scores_gemma":[0.0000022978008,0.000013423861,0.00037398838,0.0000034176485,0.000005270579,0.000013279469,0.0000026196424,0.9984578,0.00035454865,0.00064971927,0.00012015785,0.0000034298796],"about_ca_topic_score_codex":0.017269997,"about_ca_topic_score_gemma":0.017422654,"teacher_disagreement_score":0.017269997,"about_ca_system_score_codex":0.0012467562,"about_ca_system_score_gemma":0.0013707894,"threshold_uncertainty_score":0.03433895},"labels":[],"label_agreement":null},{"id":"W4406172285","doi":"10.1007/s10664-024-10555-4","title":"Lightweight dynamic build batching algorithms for continuous integration","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Process (computing); Batch processing; Quality (philosophy); Software; Quality assurance; Distributed computing; Industrial engineering; Real-time computing; Software engineering; Engineering; Operating system; Operations management","score_opus":0.014474937432760255,"score_gpt":0.3036872083430948,"score_spread":0.2892122709103346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406172285","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011334809,0.00022669432,0.9707445,0.00013401598,0.00010274383,0.000112100635,0.00017516728,0.01585379,0.001316185],"genre_scores_gemma":[0.1656081,0.00012816589,0.82667315,0.00014364648,0.0001001948,0.00025000647,0.000911056,0.0024804163,0.003705243],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956683,0.0006410976,0.00039662942,0.0009448172,0.0018368,0.0005123484],"domain_scores_gemma":[0.98669785,0.005249305,0.00074714684,0.0054043117,0.0013015588,0.0005998689],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029450704,0.0018815473,0.0019467174,0.0016594945,0.0013698693,0.002933126,0.0054858443,0.0017725712,0.015391985],"category_scores_gemma":[0.015835615,0.0018650885,0.0016568595,0.002308937,0.0015134924,0.0050861067,0.006259702,0.0042699957,0.005872539],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013890686,0.0003866611,0.0032499137,0.00026104253,0.0001469844,0.00020532458,0.00041722157,0.07761628,0.027258448,0.019783836,0.016591918,0.85269326],"study_design_scores_gemma":[0.00020864712,0.00015741478,0.0008767395,0.000030399011,0.00006815738,0.00013896496,0.00011641048,0.94441634,0.014277521,0.034064747,0.0055946936,0.00004994825],"about_ca_topic_score_codex":0.0063271946,"about_ca_topic_score_gemma":0.0117609445,"teacher_disagreement_score":0.015391985,"about_ca_system_score_codex":0.0013568883,"about_ca_system_score_gemma":0.0029560206,"threshold_uncertainty_score":0.05149126},"labels":[],"label_agreement":null},{"id":"W4406255614","doi":"10.18293/seke2024-005","title":"Comparing Machine Learning and Feature Selection Approaches for Automated Bug Report Assignment (P)","year":2024,"lang":"en","type":"article","venue":"Proceedings/Proceedings of the ... International Conference on Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Feature selection; Computer science; Artificial intelligence; Machine learning; Selection (genetic algorithm); Feature (linguistics)","score_opus":0.0359414633791206,"score_gpt":0.26875970248246533,"score_spread":0.23281823910334473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406255614","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86924076,0.0073501836,0.107633315,0.001374827,0.00059266965,0.00029818845,0.0017459545,0.006210296,0.005553651],"genre_scores_gemma":[0.9287659,0.0007456244,0.06510534,0.00017786419,0.00018956362,0.00010324102,0.0028363883,0.00015741026,0.0019187535],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99608153,0.0017793949,0.00037788815,0.00047332136,0.0009893978,0.0002985958],"domain_scores_gemma":[0.9744726,0.019772,0.0011202666,0.0013296525,0.0029086282,0.00039678148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063498183,0.00091428385,0.0008713133,0.0046995208,0.0005321471,0.0012557153,0.0010070034,0.0009885309,0.0016071146],"category_scores_gemma":[0.017405095,0.0002457005,0.0010635245,0.0024962728,0.00031778496,0.0016971624,0.0009458322,0.0007401021,0.00079368224],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026004163,0.0008098241,0.0569351,0.0004233892,0.00075167575,0.00010015965,0.00012572051,0.053324547,0.0046680756,0.0007009326,0.014252062,0.86530817],"study_design_scores_gemma":[0.00040472628,0.0018657998,0.07337095,0.00010182147,0.00061973097,0.00025788276,0.00036207488,0.9065644,0.008999551,0.0030370879,0.0043354877,0.00008047785],"about_ca_topic_score_codex":0.0072132526,"about_ca_topic_score_gemma":0.007994342,"teacher_disagreement_score":0.0072132526,"about_ca_system_score_codex":0.0007082661,"about_ca_system_score_gemma":0.001238148,"threshold_uncertainty_score":0.033581436},"labels":[],"label_agreement":null},{"id":"W4406262211","doi":"10.1109/qce60285.2024.10261","title":"Evaluating the Performance of a D-Wave Quantum Annealing System for Feature Subset Selection in Software Defect Prediction","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Feature selection; Computer science; Simulated annealing; Software; Quantum; Annealing (glass); Software system; Feature (linguistics); Artificial intelligence; Pattern recognition (psychology); Machine learning; Materials science; Physics; Operating system","score_opus":0.047914370139532515,"score_gpt":0.31862269392295595,"score_spread":0.2707083237834234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406262211","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67871803,0.00096319465,0.3131765,0.0008532126,0.00012531651,0.00019007543,0.00032051225,0.0017517407,0.0039013498],"genre_scores_gemma":[0.8390459,0.00016655796,0.15893777,0.00020195765,0.000018619614,0.00013896252,0.0005404141,0.00005689701,0.00089285464],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944884,0.00022176995,0.00004003825,0.00011041007,0.00012573524,0.000053273787],"domain_scores_gemma":[0.9971718,0.0019280018,0.00013900636,0.0002069696,0.0004602681,0.00009395376],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022192746,0.0007341862,0.00082390406,0.0007953512,0.0004808382,0.00069944485,0.0011025666,0.0012677967,0.0011279919],"category_scores_gemma":[0.005473408,0.00032775986,0.00071544823,0.00072588125,0.0005100775,0.0010255055,0.0005422134,0.00088492373,0.00022297287],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003649581,0.00024940862,0.004898176,0.00015582294,0.00010770557,0.00006511699,0.00007074695,0.91384846,0.0051097926,0.00304922,0.0014296259,0.07065101],"study_design_scores_gemma":[0.000013780126,0.000049501246,0.0002908928,0.0000022222864,0.0000056467516,0.000005178539,0.000008443423,0.99804556,0.0010909729,0.00038202934,0.000102268816,0.000003433876],"about_ca_topic_score_codex":0.007889559,"about_ca_topic_score_gemma":0.006025869,"teacher_disagreement_score":0.007889559,"about_ca_system_score_codex":0.000717285,"about_ca_system_score_gemma":0.0012972801,"threshold_uncertainty_score":0.015687287},"labels":[],"label_agreement":null},{"id":"W4406263807","doi":"10.1109/iconat61936.2024.10774589","title":"A Comprehensive Metric for Evaluating Object-Oriented Software Complexity","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hatch (Canada)","funders":"","keywords":"Computer science; Metric (unit); Software metric; Software; Software engineering; Software development; Software construction; Programming language; Engineering","score_opus":0.1039442451360682,"score_gpt":0.3767093333506895,"score_spread":0.2727650882146213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406263807","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13411656,0.0009204981,0.8447692,0.00033311488,0.00013175256,0.00053122616,0.0012366455,0.001921404,0.016039558],"genre_scores_gemma":[0.5492624,0.00041906734,0.4460084,0.00007715443,0.0000715038,0.0005304314,0.0018090118,0.00028358336,0.0015384733],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9924489,0.0011617943,0.00073930447,0.0004577434,0.005001371,0.00019082267],"domain_scores_gemma":[0.9789412,0.007089481,0.003588405,0.002317959,0.0070689237,0.0009940703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039959867,0.0009784643,0.00065240724,0.0089660315,0.0007632436,0.0023118781,0.00084671983,0.0006054548,0.0015426464],"category_scores_gemma":[0.030995427,0.00024509476,0.0005429876,0.005504529,0.0010357631,0.004109328,0.0020041182,0.00084255805,0.0004135293],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026367093,0.00040046283,0.10993107,0.0010447073,0.00043665487,0.00022159508,0.0021456042,0.04908628,0.04671139,0.074608125,0.010090988,0.70505947],"study_design_scores_gemma":[0.00009489536,0.0028806853,0.21907221,0.00067729683,0.0004564535,0.0018988749,0.003721057,0.47023907,0.06815845,0.14959262,0.08254268,0.0006657111],"about_ca_topic_score_codex":0.0016642753,"about_ca_topic_score_gemma":0.0023818077,"teacher_disagreement_score":0.0089660315,"about_ca_system_score_codex":0.0011473289,"about_ca_system_score_gemma":0.0014328341,"threshold_uncertainty_score":0.021133065},"labels":[],"label_agreement":null},{"id":"W4406421069","doi":"10.1007/s10489-024-06087-5","title":"Cross-project defect prediction based on autoencoder with dynamic adversarial adaptation","year":2025,"lang":"en","type":"article","venue":"Applied Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Autoencoder; Adversarial system; Adaptation (eye); Artificial intelligence; Machine learning; Artificial neural network; Neuroscience","score_opus":0.015150970826948402,"score_gpt":0.2902956417986149,"score_spread":0.2751446709716665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406421069","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07699007,0.00045544608,0.91948664,0.00023614804,0.00011607181,0.000044289944,0.00007662066,0.0012122559,0.0013824481],"genre_scores_gemma":[0.90058714,0.00023008724,0.09458025,0.0001749405,0.000064855834,0.00006310598,0.0003308433,0.00010172989,0.003867039],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993235,0.000115939445,0.000033424807,0.00020446864,0.00022866318,0.0000940383],"domain_scores_gemma":[0.99820566,0.00065657456,0.0001794179,0.00026137836,0.00061019225,0.00008669979],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013068713,0.0009427813,0.0010119837,0.00079224084,0.00027122415,0.0005122733,0.0012065592,0.0010714475,0.0011887918],"category_scores_gemma":[0.0031223465,0.00033323813,0.00055856595,0.00058401196,0.0005211095,0.0011496788,0.0012149186,0.0012833079,0.0004169899],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029639024,0.00023019449,0.00483419,0.00006715859,0.00013407561,0.00021788497,0.000053912623,0.6932244,0.010217709,0.0027615754,0.0032754308,0.28468704],"study_design_scores_gemma":[0.0000019504157,0.00001663501,0.00030689995,0.000002339229,0.0000064060014,0.000022205357,0.0000023472962,0.99815696,0.0009884448,0.00041661263,0.00007583372,0.0000033728552],"about_ca_topic_score_codex":0.003565435,"about_ca_topic_score_gemma":0.0038604871,"teacher_disagreement_score":0.003565435,"about_ca_system_score_codex":0.00041702748,"about_ca_system_score_gemma":0.0007086853,"threshold_uncertainty_score":0.0070893764},"labels":[],"label_agreement":null},{"id":"W4406482123","doi":"10.21203/rs.3.rs-5830055/v1","title":"Automatic Assistance to Mitigate Rollback Inconsistencies in Collaborative Edits","year":2025,"lang":"en","type":"preprint","venue":"Research Square","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Global Institute for Water Security, University of Saskatchewan","keywords":"Rollback; Computer science; Data science; World Wide Web; Database; Database transaction","score_opus":0.03737389197065287,"score_gpt":0.3769926381830503,"score_spread":0.3396187462123974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406482123","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.120881744,0.00080951327,0.7697601,0.0010377617,0.00085083867,0.00043367062,0.0011460194,0.0992494,0.00583107],"genre_scores_gemma":[0.5323957,0.00023193345,0.45038167,0.00038123986,0.00027469473,0.00015358173,0.0025908514,0.005013278,0.008577075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9864093,0.0032268243,0.0010653362,0.0028433693,0.005799486,0.00065571297],"domain_scores_gemma":[0.92837924,0.034559608,0.0056138905,0.017785687,0.012148795,0.0015127741],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006412188,0.0019423364,0.0015862468,0.0026755345,0.0015016499,0.0033939113,0.004454419,0.0035314937,0.006850753],"category_scores_gemma":[0.065569505,0.0010977509,0.0007231464,0.0018021102,0.0008878773,0.0044748667,0.0055038217,0.0028797416,0.004227575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001909469,0.0006939813,0.010931078,0.00080067076,0.00016242889,0.0014954399,0.002853867,0.0142732635,0.056358553,0.005627226,0.039394174,0.86549973],"study_design_scores_gemma":[0.0003373428,0.0006587279,0.007173065,0.00025933387,0.00033870415,0.0024935962,0.00180509,0.7460246,0.16620198,0.017963842,0.056484964,0.00025877802],"about_ca_topic_score_codex":0.002212042,"about_ca_topic_score_gemma":0.0032151332,"teacher_disagreement_score":0.006850753,"about_ca_system_score_codex":0.0004993715,"about_ca_system_score_gemma":0.0020542606,"threshold_uncertainty_score":0.033911347},"labels":[],"label_agreement":null},{"id":"W4406500031","doi":"10.1109/cascon62161.2024.10838142","title":"Translating Formal Specs: Event-B to English","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Event (particle physics); Formal methods; Programming language; Physics","score_opus":0.014997561924364357,"score_gpt":0.2778480064339703,"score_spread":0.26285044450960593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406500031","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08398652,0.00026868054,0.7715312,0.00187635,0.0006856811,0.0012166717,0.025696466,0.064071365,0.050667115],"genre_scores_gemma":[0.3576196,0.0004947704,0.56971395,0.0010865862,0.000120424986,0.0008431713,0.03977093,0.014032052,0.01631838],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99808973,0.000659667,0.0003560103,0.00032311774,0.0004521894,0.000119376295],"domain_scores_gemma":[0.99261135,0.0037845352,0.00039804727,0.0011169018,0.0019598478,0.00012926485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024203127,0.0012045422,0.0004634466,0.0007865607,0.00053926115,0.0022325676,0.0011519777,0.0007309003,0.02096151],"category_scores_gemma":[0.013307687,0.00059384597,0.00057949463,0.0006869291,0.00092563307,0.0028763507,0.0019953325,0.0013174703,0.00853404],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016976146,0.0008890849,0.0065473756,0.0038109988,0.0000958169,0.0020888427,0.01660555,0.021459872,0.10818101,0.20920235,0.20342344,0.42599797],"study_design_scores_gemma":[0.00037558103,0.0005300225,0.0039860033,0.0007128853,0.00009643575,0.0019006557,0.0051232646,0.1039724,0.14642449,0.08044501,0.6561883,0.00024498842],"about_ca_topic_score_codex":0.0041599027,"about_ca_topic_score_gemma":0.0037998436,"teacher_disagreement_score":0.02096151,"about_ca_system_score_codex":0.001030064,"about_ca_system_score_gemma":0.0017147069,"threshold_uncertainty_score":0.070123255},"labels":[],"label_agreement":null},{"id":"W4406738291","doi":"10.1145/3715006","title":"The Good, the Bad, and the Monstrous: Predicting Highly Change-Prone Source Code Methods at Their Inception","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Source code; Code (set theory); Computer security; Programming language","score_opus":0.0632421140070817,"score_gpt":0.3331406067424445,"score_spread":0.2698984927353628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406738291","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97677386,0.00079123233,0.0199747,0.00064541376,0.000024925585,0.000029990131,0.00031816875,0.00033037964,0.0011112888],"genre_scores_gemma":[0.9859098,0.00016247043,0.012746798,0.0000850195,0.000019367848,0.000014339149,0.0005178632,0.000050199647,0.00049405],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99747616,0.00073787605,0.00016362558,0.000539363,0.00077636295,0.00030660554],"domain_scores_gemma":[0.96412593,0.024062086,0.0045519844,0.0022662014,0.0035685417,0.0014252681],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006565736,0.0008637313,0.00041772355,0.004637455,0.00085097505,0.0018993593,0.0007289036,0.0012063852,0.0004527418],"category_scores_gemma":[0.029815756,0.0003996367,0.00044006205,0.0019107797,0.0011174685,0.0026339428,0.0012556394,0.0019458095,0.0004137377],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032192483,0.0003346978,0.8506148,0.00011776178,0.000079904035,0.00030252337,0.001116201,0.023595445,0.0024043343,0.0013110742,0.0024951717,0.11730613],"study_design_scores_gemma":[0.000039882772,0.00034853804,0.37988216,0.00019120285,0.00014105812,0.000684537,0.00243406,0.58978695,0.0078034136,0.013159353,0.0054398444,0.000089072826],"about_ca_topic_score_codex":0.010300768,"about_ca_topic_score_gemma":0.018761799,"teacher_disagreement_score":0.010300768,"about_ca_system_score_codex":0.00088716904,"about_ca_system_score_gemma":0.0018274263,"threshold_uncertainty_score":0.03472334},"labels":[],"label_agreement":null},{"id":"W4406745366","doi":"10.1145/3757912","title":"Build Optimization: A Systematic Literature Review","year":2025,"lang":"en","type":"preprint","venue":"ACM Computing Surveys","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs","keywords":"Systematic review; Computer science; Management science; Political science; Engineering; MEDLINE","score_opus":0.024913450470017976,"score_gpt":0.30787717498273887,"score_spread":0.2829637245127209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406745366","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013559535,0.99292386,0.00063051016,0.00075558363,0.00014053678,0.0005336259,0.002752014,0.000033187833,0.00087469834],"genre_scores_gemma":[0.007480589,0.9867329,0.002099567,0.0007836566,0.00009416277,0.0010403588,0.0015780168,0.000020748392,0.00016994975],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.99057627,0.0025708699,0.0037540488,0.0009284951,0.001879722,0.00029061304],"domain_scores_gemma":[0.92926955,0.053536944,0.008337785,0.001209861,0.0069611412,0.000684774],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010880947,0.001948943,0.005685732,0.029500123,0.0009738025,0.003455822,0.003043458,0.002232283,0.008954866],"category_scores_gemma":[0.06536649,0.0014806582,0.006778144,0.02710449,0.0009466296,0.004447328,0.002618605,0.0016880463,0.0011702295],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016447305,0.000040659415,0.0011590287,0.88908005,0.0029486327,0.00014003742,0.00031128444,0.0002848699,0.00017848988,0.0006593915,0.0076582506,0.09737486],"study_design_scores_gemma":[0.00011733658,0.00013399989,0.0036,0.93024355,0.013755287,0.0002868066,0.00046676453,0.00011804817,0.0001821012,0.00068736624,0.050359644,0.00004904362],"about_ca_topic_score_codex":0.0113189025,"about_ca_topic_score_gemma":0.037207104,"teacher_disagreement_score":0.98911905,"about_ca_system_score_codex":0.0045687263,"about_ca_system_score_gemma":0.024488432,"threshold_uncertainty_score":0.05754465},"labels":[],"label_agreement":null},{"id":"W4406803861","doi":"10.1145/3689187.3709615","title":"Introducing Code Quality at CS1 Level: Examples and Activities","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Programming language; Software quality; Quality (philosophy); Code (set theory); Software engineering; Software; Software development","score_opus":0.06552727946841341,"score_gpt":0.33555063341636104,"score_spread":0.27002335394794763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406803861","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75314385,0.0027239989,0.1470366,0.0041849976,0.00023510717,0.0006778756,0.00082582614,0.0046364986,0.08653524],"genre_scores_gemma":[0.8415723,0.0019684604,0.13791439,0.00042762162,0.000059584156,0.00026466307,0.0010682893,0.00075758866,0.015967146],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99616927,0.0017118155,0.00020038713,0.00029677385,0.0012298471,0.00039196722],"domain_scores_gemma":[0.97219026,0.017764585,0.0014895923,0.0021357925,0.004318764,0.002100924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031563924,0.00073848764,0.0002671334,0.002357124,0.0012825992,0.0016347723,0.0012548675,0.0012165948,0.0035884497],"category_scores_gemma":[0.018644447,0.00025370583,0.0004008095,0.0019331898,0.001274049,0.0016486903,0.002286329,0.0018129495,0.001660355],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083875685,0.0029621965,0.035760485,0.0037961085,0.0000588788,0.0040039397,0.10460317,0.010811055,0.0303265,0.02928896,0.043093257,0.73445666],"study_design_scores_gemma":[0.00031561474,0.00400895,0.13278297,0.004798678,0.00015491674,0.010334837,0.04966691,0.047621474,0.082882896,0.033435185,0.6335951,0.0004024835],"about_ca_topic_score_codex":0.0031689466,"about_ca_topic_score_gemma":0.006083689,"teacher_disagreement_score":0.0035884497,"about_ca_system_score_codex":0.0015956546,"about_ca_system_score_gemma":0.0011034701,"threshold_uncertainty_score":0.016692817},"labels":[],"label_agreement":null},{"id":"W4406858280","doi":"10.1109/tse.2025.3534027","title":"Recovering Traceability Links Between Code and Documentation: A Retrospective","year":2025,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Traceability; Documentation; Computer science; Code (set theory); Programming language; Software engineering","score_opus":0.011972415143162336,"score_gpt":0.26600079791641557,"score_spread":0.25402838277325324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406858280","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4956492,0.2109145,0.19591916,0.018602502,0.0027169255,0.0010610844,0.009593144,0.0029889059,0.06255454],"genre_scores_gemma":[0.6921091,0.11041187,0.14885843,0.0030844104,0.0012442343,0.0005223668,0.01417673,0.0021460808,0.027446726],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9832908,0.0049704425,0.002133326,0.0027862585,0.0063447184,0.00047437247],"domain_scores_gemma":[0.77048224,0.09430622,0.01879785,0.032174177,0.08202225,0.0022172327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028283417,0.0007125455,0.00068506587,0.017758086,0.0020590709,0.0064904615,0.0017297675,0.0014219654,0.003059176],"category_scores_gemma":[0.1475353,0.0011739128,0.0006358514,0.011864676,0.005031139,0.009780777,0.0035896997,0.0039115823,0.0026608696],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000417273,0.0005323506,0.0407283,0.0017484303,0.00013760202,0.00041007533,0.0091915,0.001733067,0.0046542464,0.020186165,0.016810497,0.9034505],"study_design_scores_gemma":[0.000066398155,0.0012925798,0.09940371,0.0065810727,0.0004430494,0.0028037648,0.007398463,0.005175687,0.04463681,0.016609164,0.8152593,0.00033008048],"about_ca_topic_score_codex":0.0109115,"about_ca_topic_score_gemma":0.009553482,"teacher_disagreement_score":0.028283417,"about_ca_system_score_codex":0.0037814246,"about_ca_system_score_gemma":0.006361611,"threshold_uncertainty_score":0.14957881},"labels":[],"label_agreement":null},{"id":"W4407031081","doi":"10.32388/8g8tb2","title":"Enhancing Code LLMs with Reinforcement Learning in Code Generation: A Survey","year":2025,"lang":"en","type":"preprint","venue":"Qeios","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Reinforcement learning; Computer science; Compiler; Code generation; Code (set theory); Resource allocation; Resource (disambiguation); Artificial intelligence; Programming language; Computer security","score_opus":0.04221355533285516,"score_gpt":0.2985258908404057,"score_spread":0.25631233550755056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407031081","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021704057,0.096609965,0.85963166,0.0020239234,0.00018162507,0.00022423407,0.00018640108,0.003899074,0.015539079],"genre_scores_gemma":[0.21491458,0.10041814,0.6734844,0.00093041343,0.0003758024,0.00040938193,0.0008108813,0.0019046555,0.0067517688],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998243,0.0005449657,0.00014801692,0.00031005815,0.00066490803,0.00008916842],"domain_scores_gemma":[0.99383694,0.0043431944,0.0002531305,0.0007863471,0.0006775526,0.000102752245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025722685,0.0009917638,0.00078904675,0.0015356635,0.00036703594,0.001279756,0.0017865535,0.0010523445,0.002544339],"category_scores_gemma":[0.010374578,0.0006974357,0.00089348387,0.0016768433,0.0011751478,0.0017844577,0.0012160977,0.0015647358,0.0014376171],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055098742,0.0001401338,0.002037632,0.0019660674,0.000050780694,0.000058757283,0.00022673451,0.07655607,0.005563887,0.024312641,0.003949814,0.88508236],"study_design_scores_gemma":[0.00008675765,0.00063743116,0.0023583346,0.0016592958,0.00012073054,0.0007633731,0.00028906288,0.68072987,0.03223961,0.06851258,0.2124731,0.00012991403],"about_ca_topic_score_codex":0.0020030607,"about_ca_topic_score_gemma":0.0016237237,"teacher_disagreement_score":0.0025722685,"about_ca_system_score_codex":0.00088533823,"about_ca_system_score_gemma":0.0016452083,"threshold_uncertainty_score":0.013603628},"labels":[],"label_agreement":null},{"id":"W4407057502","doi":"10.1007/978-3-658-45877-5_4","title":"Knowledge Graph Based Hard Drive Failure Prediction","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Graph; Forensic engineering; Engineering; Theoretical computer science","score_opus":0.016203832910941945,"score_gpt":0.23970192886099972,"score_spread":0.2234980959500578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407057502","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032458812,0.0033690215,0.92846376,0.000699576,0.00038817438,0.000093485985,0.0021754298,0.0053084884,0.027043266],"genre_scores_gemma":[0.6974635,0.0038007046,0.22906613,0.00043721945,0.00020155055,0.000107322354,0.00521208,0.0004825296,0.06322894],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998253,0.000017075881,0.0000059686313,0.00004660966,0.000086607644,0.000018429848],"domain_scores_gemma":[0.999509,0.00028035446,0.000031898147,0.00006524684,0.00009748042,0.000016072301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00021551808,0.0007046218,0.0007476584,0.0011167566,0.00026052602,0.00070402876,0.0012269199,0.0006287876,0.008846711],"category_scores_gemma":[0.0010986269,0.00023182119,0.00054134947,0.0010278677,0.0002596826,0.0012082631,0.00052928826,0.0008079223,0.002318136],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009107175,0.000120939854,0.0012918952,0.00014375885,0.000071841285,0.00012934774,0.00002638247,0.40251708,0.0029661849,0.008111936,0.019413274,0.5651163],"study_design_scores_gemma":[0.0000036044053,0.000027375969,0.0010414219,0.000024256908,0.000022943432,0.00006702154,0.000015260637,0.97730756,0.0015943678,0.01660914,0.0032774494,0.00000968185],"about_ca_topic_score_codex":0.011263896,"about_ca_topic_score_gemma":0.013058851,"teacher_disagreement_score":0.011263896,"about_ca_system_score_codex":0.00060169084,"about_ca_system_score_gemma":0.0004475157,"threshold_uncertainty_score":0.029595137},"labels":[],"label_agreement":null},{"id":"W4407241478","doi":"10.1145/3711903","title":"Obfuscated Clone Search in JavaScript based on Reinforcement Subsequence Learning","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada); Defence Research and Development Canada; McGill University; Queen's University","funders":"","keywords":"Computer science; JavaScript; Scripting language; Codebase; Code refactoring; Programming language; Reinforcement learning; Artificial intelligence; Automatic summarization; Source code; Code (set theory); Theoretical computer science; Software","score_opus":0.0702671741606317,"score_gpt":0.32818521897802067,"score_spread":0.25791804481738895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407241478","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46735638,0.0013903937,0.51761335,0.00056531525,0.00011312043,0.00018664516,0.000151685,0.009867871,0.0027551649],"genre_scores_gemma":[0.89482176,0.00012598878,0.10243202,0.00023201694,0.00003048908,0.00006415599,0.0002842991,0.00017561074,0.0018335246],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993325,0.00016419248,0.00004090408,0.00019891218,0.000175112,0.00008839183],"domain_scores_gemma":[0.9976501,0.0012595352,0.00027668182,0.0003334929,0.00034524038,0.00013500094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011521956,0.00084620283,0.0010632003,0.0006333886,0.00036224697,0.00043723747,0.0013259231,0.0011576959,0.00077803415],"category_scores_gemma":[0.004480189,0.00030928783,0.00056102534,0.00041856695,0.00080450246,0.0012565372,0.0008156558,0.001250961,0.0003378868],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040170387,0.00037387968,0.006999155,0.00013329119,0.00007623716,0.00028544705,0.00016413057,0.6346679,0.016220888,0.0028004686,0.0021402773,0.33573663],"study_design_scores_gemma":[0.000011780076,0.00007798993,0.00027216255,0.000004027318,0.000007020144,0.00003491526,0.000008192811,0.9961994,0.0021780543,0.0010178447,0.00018409749,0.0000043625914],"about_ca_topic_score_codex":0.006831211,"about_ca_topic_score_gemma":0.007364699,"teacher_disagreement_score":0.006831211,"about_ca_system_score_codex":0.0010222932,"about_ca_system_score_gemma":0.0012409065,"threshold_uncertainty_score":0.013582885},"labels":[],"label_agreement":null},{"id":"W4407277384","doi":"10.1007/s10664-024-10609-7","title":"UPC sentinel: An accurate approach for detecting upgradeability proxy contracts in Ethereum","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Proxy (statistics); Computer science; Engineering; Machine learning","score_opus":0.03488539790364532,"score_gpt":0.326717984410252,"score_spread":0.2918325865066067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407277384","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45634282,0.0010784997,0.4988465,0.00084359146,0.0002595048,0.000492163,0.008121933,0.020821372,0.013193676],"genre_scores_gemma":[0.82896894,0.00020170229,0.16118024,0.00024341267,0.000058049438,0.00015795344,0.0049654525,0.0005309244,0.003693363],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99132735,0.0015160501,0.0006259823,0.0010701393,0.0048485873,0.00061183673],"domain_scores_gemma":[0.97545683,0.009571331,0.004144902,0.0052281497,0.0048184,0.00078034785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068490854,0.00086470874,0.0007821155,0.006083718,0.0006833981,0.0019618436,0.001993849,0.002215171,0.0035052253],"category_scores_gemma":[0.033676255,0.00042874538,0.0004697494,0.0026019178,0.00063911953,0.003951066,0.0022965597,0.0017877596,0.0012257192],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014474195,0.00075638306,0.3178537,0.000662634,0.00031954265,0.0016375124,0.0009296052,0.06455966,0.047510244,0.043618828,0.034198567,0.48650596],"study_design_scores_gemma":[0.00011347872,0.00032127523,0.05063699,0.00016065914,0.00011355873,0.00096218084,0.0004933604,0.8448142,0.05025655,0.02795029,0.024051916,0.00012552722],"about_ca_topic_score_codex":0.005484115,"about_ca_topic_score_gemma":0.007028841,"teacher_disagreement_score":0.0068490854,"about_ca_system_score_codex":0.0010829062,"about_ca_system_score_gemma":0.0024403029,"threshold_uncertainty_score":0.036221802},"labels":[],"label_agreement":null},{"id":"W4407685505","doi":"10.1016/j.scico.2025.103284","title":"Graph neural network-based long method and blob code smell detection","year":2025,"lang":"en","type":"article","venue":"Science of Computer Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Graph; Code (set theory); Artificial intelligence; Artificial neural network; Theoretical computer science; Programming language","score_opus":0.013471483309353492,"score_gpt":0.29774889675374216,"score_spread":0.28427741344438867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407685505","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07347879,0.0012671108,0.91465116,0.0006704042,0.00022418838,0.00026730742,0.0007498281,0.005841559,0.0028497055],"genre_scores_gemma":[0.67597944,0.0005857301,0.31101373,0.0006273583,0.00013115503,0.00034344257,0.0025929806,0.00060620403,0.00811986],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99852604,0.0001996369,0.000090761045,0.00056755636,0.00045071365,0.00016543207],"domain_scores_gemma":[0.9969608,0.00088360783,0.00053759356,0.00044692055,0.0009923019,0.00017884656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012763774,0.0014495468,0.0010160325,0.0029901708,0.0005214777,0.0013060773,0.002035892,0.0016379997,0.0016422208],"category_scores_gemma":[0.005140391,0.00038952995,0.0010540984,0.0017347406,0.0008456236,0.0023725778,0.0011508222,0.0017076285,0.00080163655],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007398983,0.00038322067,0.021411724,0.00034538034,0.0002542953,0.0005070721,0.00018601041,0.18728723,0.022775704,0.0074581313,0.012783357,0.74586797],"study_design_scores_gemma":[0.00000823584,0.00004940048,0.0017469153,0.000013595592,0.000018972036,0.00007447181,0.000016628828,0.99093086,0.0037926747,0.0024100465,0.0009231228,0.000015116165],"about_ca_topic_score_codex":0.013350663,"about_ca_topic_score_gemma":0.017670806,"teacher_disagreement_score":0.013350663,"about_ca_system_score_codex":0.0015764274,"about_ca_system_score_gemma":0.0013135474,"threshold_uncertainty_score":0.026545942},"labels":[],"label_agreement":null},{"id":"W4407986476","doi":"10.1007/s10664-025-10618-0","title":"Evaluating interactive documentation for programmers","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Computer science; World Wide Web; Software engineering; Programming language","score_opus":0.04750618481790624,"score_gpt":0.41323817976843863,"score_spread":0.3657319949505324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407986476","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9755727,0.0005164751,0.014571881,0.00021075118,0.000045740475,0.00018781425,0.0001602158,0.0015912196,0.0071433657],"genre_scores_gemma":[0.97256136,0.00019045892,0.023962332,0.000061846105,0.000029138877,0.000105446554,0.0005075435,0.00025688164,0.002325118],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98577785,0.007269369,0.0010237555,0.00081282074,0.0046645245,0.00045162116],"domain_scores_gemma":[0.7851467,0.15683575,0.013870568,0.014107821,0.02555297,0.004486238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010161582,0.0009799625,0.00084742083,0.0022913225,0.0011043303,0.0037444402,0.001555256,0.0019670967,0.004311485],"category_scores_gemma":[0.16134547,0.00047942656,0.00044915397,0.0016756401,0.00069476257,0.0029797165,0.0020344597,0.0012407133,0.0009837829],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008655415,0.007034823,0.12522174,0.0014286612,0.0004671304,0.00070318184,0.00846014,0.037291843,0.018828016,0.0032410948,0.006501003,0.7821669],"study_design_scores_gemma":[0.0034786016,0.036334567,0.2538049,0.0015982052,0.0019740395,0.0016991849,0.013093629,0.5502153,0.09596464,0.0146487495,0.026720934,0.00046730886],"about_ca_topic_score_codex":0.004453322,"about_ca_topic_score_gemma":0.0061037233,"teacher_disagreement_score":0.010161582,"about_ca_system_score_codex":0.0016472593,"about_ca_system_score_gemma":0.0024823935,"threshold_uncertainty_score":0.053740203},"labels":[],"label_agreement":null},{"id":"W4408014589","doi":"10.22152/programming-journal.org/2026/10/8","title":"Dynamic Program Slices Change How Developers Diagnose Gradual Run-Time Type Errors","year":2025,"lang":"en","type":"article","venue":"The Art Science and Engineering of Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Alberta","funders":"","keywords":"Computer science; Type (biology); Programming language; Geology","score_opus":0.019183461648924895,"score_gpt":0.27848719916219394,"score_spread":0.2593037375132691,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408014589","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20707086,0.00052014145,0.73605907,0.0018718404,0.00032872302,0.0005070352,0.00047914384,0.04549514,0.007668056],"genre_scores_gemma":[0.54214054,0.00029619262,0.44906223,0.00067622337,0.000049688482,0.00021294245,0.00043962526,0.0045235003,0.0025990834],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9931706,0.0024655077,0.0004577926,0.0012857134,0.0020478975,0.00057240337],"domain_scores_gemma":[0.9566626,0.017626258,0.0038433713,0.014823453,0.005624374,0.0014199116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007649836,0.0014592168,0.0006654743,0.0016052735,0.0008764012,0.003002524,0.0025439928,0.0020213278,0.0034435864],"category_scores_gemma":[0.068720914,0.0016502199,0.00096032536,0.00060079544,0.0027100286,0.008271688,0.0035675399,0.0029169868,0.0011757425],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019681742,0.0006962153,0.08490436,0.0011583185,0.00030754824,0.0042310534,0.02735409,0.04799831,0.16200997,0.060915682,0.021470472,0.58698577],"study_design_scores_gemma":[0.0005103945,0.002174337,0.029713303,0.0016977786,0.00069523853,0.004637073,0.00821431,0.31748316,0.35396796,0.11889274,0.16105722,0.0009564857],"about_ca_topic_score_codex":0.0029847228,"about_ca_topic_score_gemma":0.0039365645,"teacher_disagreement_score":0.007649836,"about_ca_system_score_codex":0.00081137294,"about_ca_system_score_gemma":0.0031145227,"threshold_uncertainty_score":0.040456712},"labels":[],"label_agreement":null},{"id":"W4408054024","doi":"10.1145/3721125","title":"Unraveling Code Clone Dynamics in Deep Learning Frameworks","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Toronto; Université du Québec à Montréal","funders":"","keywords":"Computer science; Code (set theory); Software engineering; Artificial intelligence; Data science; Programming language","score_opus":0.03839091002517494,"score_gpt":0.32376943785259454,"score_spread":0.2853785278274196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408054024","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98729557,0.00051537843,0.010925586,0.00025589645,0.000007378844,0.000017075201,0.00008984713,0.00031401808,0.00057926914],"genre_scores_gemma":[0.993064,0.000098761324,0.006292526,0.000045632576,0.0000043997125,0.000019204734,0.00018916017,0.000058479956,0.00022791955],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9972941,0.0006618699,0.00016535439,0.0006855628,0.0008299243,0.00036305442],"domain_scores_gemma":[0.9718847,0.014176072,0.0070490926,0.0027823672,0.0032569459,0.00085070584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045671156,0.0004445414,0.00043392842,0.0030204481,0.0006622813,0.001526167,0.0011812065,0.0007431563,0.0004457586],"category_scores_gemma":[0.040502697,0.0003634454,0.00047388574,0.0018757641,0.0017299923,0.0039462578,0.001686953,0.001324553,0.000114804345],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027370747,0.00025172537,0.77916527,0.00020349058,0.00017752252,0.00071214227,0.0051459707,0.044888817,0.007687228,0.006488162,0.0014714693,0.15353449],"study_design_scores_gemma":[0.000035651396,0.00027990656,0.35597098,0.00012185644,0.00013580605,0.0007790343,0.0024781493,0.6141773,0.0061908187,0.014515616,0.005236694,0.00007812979],"about_ca_topic_score_codex":0.010689818,"about_ca_topic_score_gemma":0.014433778,"teacher_disagreement_score":0.010689818,"about_ca_system_score_codex":0.0017663075,"about_ca_system_score_gemma":0.001342619,"threshold_uncertainty_score":0.02415353},"labels":[],"label_agreement":null},{"id":"W4408069070","doi":"10.1007/s42484-025-00236-w","title":"Comparative analysis of quantum and classical support vector classifiers for software bug prediction: an exploratory study","year":2025,"lang":"en","type":"article","venue":"Quantum Machine Intelligence","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Prince Edward Island; University of Saskatchewan","funders":"","keywords":"Support vector machine; Computer science; Quantum; Software; Exploratory research; Artificial intelligence; Machine learning; Physics; Programming language; Quantum mechanics; Sociology","score_opus":0.05253343009886388,"score_gpt":0.35429137319863224,"score_spread":0.30175794309976833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408069070","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9590218,0.0023936417,0.034626517,0.00038666543,0.000055003115,0.000055820237,0.00024562873,0.00021548352,0.0029993805],"genre_scores_gemma":[0.99170285,0.00028516387,0.0073557813,0.000031045995,0.000029617659,0.000016106034,0.00023456874,0.000015732881,0.00032911304],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968925,0.0012699218,0.00018350752,0.0003876119,0.0010988349,0.00016776312],"domain_scores_gemma":[0.94981116,0.042508997,0.001327776,0.0013442668,0.004670994,0.0003366666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073502404,0.00038594566,0.00076785556,0.0019599136,0.00046038243,0.0012685667,0.0008973005,0.00070874684,0.001408462],"category_scores_gemma":[0.033678077,0.00013258753,0.000498702,0.0016359268,0.0005313952,0.0030481333,0.0005464366,0.0005167796,0.00027427185],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004879621,0.001958658,0.16241966,0.0010235916,0.0010086126,0.00023974144,0.0007186189,0.09750018,0.007404196,0.015359557,0.004903529,0.7025841],"study_design_scores_gemma":[0.00011201945,0.0017965783,0.047445606,0.0000720071,0.00035912433,0.00021742124,0.0007276663,0.9355704,0.0029514083,0.009076138,0.0016276173,0.00004400857],"about_ca_topic_score_codex":0.0023489604,"about_ca_topic_score_gemma":0.001741655,"teacher_disagreement_score":0.0073502404,"about_ca_system_score_codex":0.0007749751,"about_ca_system_score_gemma":0.0007380845,"threshold_uncertainty_score":0.038872242},"labels":[],"label_agreement":null},{"id":"W4408069583","doi":"10.1007/s10664-025-10616-2","title":"WIA-SZZ: Work item aware SZZ","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Work (physics); Engineering; Mechanical engineering","score_opus":0.0177152079228471,"score_gpt":0.2842851256222398,"score_spread":0.2665699176993927,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408069583","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032600593,0.00058282336,0.33466172,0.00074152707,0.0004192902,0.0015785985,0.09026908,0.5134727,0.025673741],"genre_scores_gemma":[0.17383416,0.00058131467,0.5981326,0.00064536306,0.0001556779,0.0024710929,0.16390324,0.018196555,0.042079955],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982437,0.0002319233,0.00017969291,0.0004419606,0.0007207162,0.00018205504],"domain_scores_gemma":[0.9968045,0.0006631133,0.00024334999,0.0014168209,0.0006818901,0.00019034858],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001287796,0.0018974203,0.001076189,0.0031634879,0.0004992278,0.0021899617,0.0025461572,0.0008611985,0.028726475],"category_scores_gemma":[0.008261156,0.0010043273,0.0015085356,0.0021420838,0.00030420642,0.0035229044,0.004710234,0.0013883184,0.025143735],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023230405,0.00083840993,0.013849091,0.0011083082,0.00050981226,0.00019669368,0.00036292232,0.00480368,0.017973572,0.007623559,0.3149212,0.6354897],"study_design_scores_gemma":[0.0016364404,0.0009082046,0.04460111,0.00031885633,0.00075105735,0.00073606713,0.00068515923,0.3154838,0.10613168,0.04901392,0.47917873,0.00055497367],"about_ca_topic_score_codex":0.0028458268,"about_ca_topic_score_gemma":0.0062377728,"teacher_disagreement_score":0.028726475,"about_ca_system_score_codex":0.00042544494,"about_ca_system_score_gemma":0.0011361338,"threshold_uncertainty_score":0.096099675},"labels":[],"label_agreement":null},{"id":"W4408089088","doi":"10.1007/s10664-025-10634-0","title":"Fixer-level supervised contrastive learning for bug assignment","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Artificial intelligence; Natural language processing; Machine learning","score_opus":0.03540607240363587,"score_gpt":0.3021084142661999,"score_spread":0.26670234186256403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408089088","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11997999,0.0011610042,0.8658953,0.0004143509,0.00016909579,0.00016531044,0.00072324136,0.008120787,0.0033708264],"genre_scores_gemma":[0.7564338,0.00015507427,0.23550095,0.00032640854,0.0001547782,0.00015613383,0.002598404,0.0005666358,0.0041077468],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99757904,0.00083586416,0.00011202119,0.0008624039,0.00042112079,0.0001895583],"domain_scores_gemma":[0.98897564,0.0070190323,0.00063052744,0.0016155108,0.0014482643,0.00031106337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035625426,0.0011288207,0.0014197892,0.0020121366,0.0008094157,0.0011499039,0.0038392125,0.002226822,0.0033674287],"category_scores_gemma":[0.014094843,0.00044876433,0.0008362628,0.0011876198,0.00093292625,0.0024527677,0.0023738998,0.0032878781,0.0011872584],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014241136,0.00097499933,0.012149215,0.00042437084,0.000314726,0.00016613229,0.00020014418,0.091114506,0.022100085,0.008111721,0.013056705,0.84996337],"study_design_scores_gemma":[0.00007083501,0.0002443716,0.0020267242,0.000023880099,0.00005382218,0.00009129818,0.000028914827,0.98169655,0.0060343584,0.008221787,0.0014873053,0.000020114105],"about_ca_topic_score_codex":0.0034988553,"about_ca_topic_score_gemma":0.00890047,"teacher_disagreement_score":0.0038392125,"about_ca_system_score_codex":0.0011540448,"about_ca_system_score_gemma":0.0016191461,"threshold_uncertainty_score":0.01884073},"labels":[],"label_agreement":null},{"id":"W4408106796","doi":"10.22152/programming-journal.org/2025/10/8","title":"Dynamic Program Slices Change How Developers Diagnose Gradual Run-Time Type Errors","year":2025,"lang":"en","type":"article","venue":"The Art Science and Engineering of Programming","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Alberta","funders":"","keywords":"Computer science; Programming language; Parallel computing","score_opus":0.019183461648924895,"score_gpt":0.27848719916219394,"score_spread":0.2593037375132691,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408106796","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20707086,0.00052014145,0.73605907,0.0018718404,0.00032872302,0.0005070352,0.00047914384,0.04549514,0.007668056],"genre_scores_gemma":[0.54214054,0.00029619262,0.44906223,0.00067622337,0.000049688482,0.00021294245,0.00043962526,0.0045235003,0.0025990834],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9931706,0.0024655077,0.0004577926,0.0012857134,0.0020478975,0.00057240337],"domain_scores_gemma":[0.9566626,0.017626258,0.0038433713,0.014823453,0.005624374,0.0014199116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007649836,0.0014592168,0.0006654743,0.0016052735,0.0008764012,0.003002524,0.0025439928,0.0020213278,0.0034435864],"category_scores_gemma":[0.068720914,0.0016502199,0.00096032536,0.00060079544,0.0027100286,0.008271688,0.0035675399,0.0029169868,0.0011757425],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019681742,0.0006962153,0.08490436,0.0011583185,0.00030754824,0.0042310534,0.02735409,0.04799831,0.16200997,0.060915682,0.021470472,0.58698577],"study_design_scores_gemma":[0.0005103945,0.002174337,0.029713303,0.0016977786,0.00069523853,0.004637073,0.00821431,0.31748316,0.35396796,0.11889274,0.16105722,0.0009564857],"about_ca_topic_score_codex":0.0029847228,"about_ca_topic_score_gemma":0.0039365645,"teacher_disagreement_score":0.007649836,"about_ca_system_score_codex":0.00081137294,"about_ca_system_score_gemma":0.0031145227,"threshold_uncertainty_score":0.040456712},"labels":[],"label_agreement":null},{"id":"W4408162704","doi":"10.1007/978-3-031-73143-3_15","title":"Practical Guidelines for the Selection and Evaluation of Natural Language Processing Techniques in Requirements Engineering","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Selection (genetic algorithm); Computer science; Natural (archaeology); Natural language processing; Software engineering; Linguistics; Systems engineering; Artificial intelligence; Engineering; History; Archaeology; Philosophy","score_opus":0.10887866162592745,"score_gpt":0.43236194165666275,"score_spread":0.3234832800307353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408162704","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026281187,0.003649107,0.9273482,0.0042884536,0.0003573331,0.002595272,0.0011788289,0.007982549,0.049972247],"genre_scores_gemma":[0.005153777,0.0010933102,0.98541194,0.00036906538,0.00006130201,0.0009487319,0.0007618867,0.0006731274,0.005526911],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97655565,0.012311006,0.0028581063,0.0008662878,0.006976555,0.00043240943],"domain_scores_gemma":[0.91697437,0.053010367,0.002463711,0.006810969,0.019842537,0.0008981101],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022512773,0.0024260622,0.0014904062,0.0082077915,0.0018966467,0.0081884945,0.005358013,0.0037191121,0.030261863],"category_scores_gemma":[0.09165264,0.0018044327,0.0010766279,0.006487836,0.0019064313,0.0076580215,0.002802279,0.0038717112,0.015150359],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017686558,0.0006837005,0.00076726306,0.0019989128,0.00004364707,0.0005733438,0.0015925888,0.00584865,0.011018078,0.14656112,0.101311356,0.7294246],"study_design_scores_gemma":[0.0005211338,0.0005436104,0.002748956,0.0068557602,0.00018465212,0.0018208781,0.0027676593,0.078285545,0.033460952,0.3545016,0.5179192,0.00039007977],"about_ca_topic_score_codex":0.0029523785,"about_ca_topic_score_gemma":0.0077127474,"teacher_disagreement_score":0.030261863,"about_ca_system_score_codex":0.0017631793,"about_ca_system_score_gemma":0.0032705006,"threshold_uncertainty_score":0.1190604},"labels":[],"label_agreement":null},{"id":"W4408162729","doi":"10.1007/978-3-031-73143-3_14","title":"Empirical Evaluation of Tools for Hairy Natural Language Requirements Engineering Tasks","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Natural language; Software engineering; Linguistics; Natural language processing; Philosophy","score_opus":0.09265498278320823,"score_gpt":0.3723471088161209,"score_spread":0.27969212603291266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408162729","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9491134,0.0010184398,0.02940136,0.000370087,0.00009864925,0.00096324354,0.0019191173,0.0035819912,0.013533643],"genre_scores_gemma":[0.92527807,0.00048791632,0.060426425,0.0001961303,0.00003691609,0.00090354524,0.0054349615,0.000624423,0.006611648],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9902789,0.0060454626,0.00076245377,0.0007618912,0.0018640301,0.0002872846],"domain_scores_gemma":[0.7732897,0.20488645,0.0042292895,0.009401738,0.0065255486,0.0016671957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007880617,0.00082574645,0.000488731,0.0018595478,0.0006232558,0.0015443806,0.0020349866,0.0010908004,0.005756217],"category_scores_gemma":[0.09306299,0.00044859754,0.00043321008,0.0014011827,0.00082484673,0.0027094444,0.0027030602,0.0013736254,0.0026558936],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011446263,0.013072192,0.031525683,0.0045336084,0.0002874374,0.0009102696,0.013721774,0.017322494,0.024665069,0.005587536,0.03155995,0.84536767],"study_design_scores_gemma":[0.009008178,0.035439577,0.2734838,0.0026220768,0.0010836684,0.002985485,0.029829321,0.4390621,0.0787191,0.02021726,0.106913194,0.000636279],"about_ca_topic_score_codex":0.002921692,"about_ca_topic_score_gemma":0.0045213825,"teacher_disagreement_score":0.007880617,"about_ca_system_score_codex":0.0009933175,"about_ca_system_score_gemma":0.000995193,"threshold_uncertainty_score":0.041677177},"labels":[],"label_agreement":null},{"id":"W4408162739","doi":"10.1007/978-3-031-73143-3_5","title":"Detecting Defects in Natural Language Requirements Specifications","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sheridan College; University of Waterloo","funders":"","keywords":"Computer science; Natural (archaeology); Programming language; History; Archaeology","score_opus":0.04183263579618568,"score_gpt":0.2915067177688678,"score_spread":0.2496740819726821,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408162739","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04118354,0.0008115043,0.91628456,0.0007927365,0.00015375506,0.00021647339,0.0008421439,0.008894529,0.0308208],"genre_scores_gemma":[0.19511952,0.00097409566,0.7559483,0.0004342051,0.00006361613,0.00012481675,0.0038053717,0.0021243368,0.04140579],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99824995,0.00031902405,0.00009834212,0.00020565274,0.0010476939,0.00007933604],"domain_scores_gemma":[0.9930609,0.005111358,0.00035513862,0.0005975179,0.0008360013,0.00003900146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012504566,0.0007354508,0.00045987807,0.0014984637,0.00025247614,0.0011576689,0.0012372843,0.00095200265,0.0067663915],"category_scores_gemma":[0.009326797,0.00063413486,0.0009046428,0.00075606816,0.0005934219,0.0019200816,0.0008394772,0.0011092938,0.0029932868],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000110639754,0.00017340566,0.0048640408,0.000829974,0.00004531013,0.001181047,0.0007540987,0.01930026,0.03938987,0.050207835,0.034190707,0.8489529],"study_design_scores_gemma":[0.000064907625,0.00046439114,0.009491868,0.00096982653,0.0001378013,0.008142491,0.0010329492,0.4853353,0.15084253,0.18519555,0.15818693,0.0001354233],"about_ca_topic_score_codex":0.0014300763,"about_ca_topic_score_gemma":0.0029303853,"teacher_disagreement_score":0.0067663915,"about_ca_system_score_codex":0.00051313755,"about_ca_system_score_gemma":0.0005811415,"threshold_uncertainty_score":0.022635818},"labels":[],"label_agreement":null},{"id":"W4408256992","doi":"10.1145/3711403.3711450","title":"Automated Generation of Challenge Questions for Student Code Evaluation Using Abstract Syntax Tree Embeddings and RAG: An Exploratory Study","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Abstract syntax tree; Computer science; Programming language; Syntax; Tree (set theory); Abstract syntax; Code (set theory); Exploratory research; Code generation; Theoretical computer science; Natural language processing; Mathematics; Operating system; Sociology; Combinatorics","score_opus":0.15574449195593493,"score_gpt":0.42529055174890884,"score_spread":0.2695460597929739,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408256992","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9268885,0.00012454543,0.069031835,0.00014602707,0.00003428542,0.0010225428,0.00047366408,0.0010096006,0.001269018],"genre_scores_gemma":[0.90062135,0.000088469875,0.09537876,0.00010436544,0.000027354645,0.0011094968,0.0010996042,0.00024107526,0.0013294847],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98107713,0.013400362,0.000844829,0.0018218851,0.0022679782,0.0005877601],"domain_scores_gemma":[0.8214889,0.1474102,0.006110145,0.008453824,0.014579342,0.0019575309],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011193395,0.0012685492,0.0008828999,0.0016394655,0.0005798175,0.0018111246,0.0016689625,0.0017924117,0.0015268059],"category_scores_gemma":[0.08589006,0.0004061338,0.00061673817,0.0008724997,0.0009009131,0.0019963782,0.0023194938,0.0015544521,0.0008336953],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040591797,0.013569792,0.12382211,0.0028217717,0.000305005,0.0024518198,0.06900129,0.045411322,0.10307278,0.0046201665,0.009430264,0.6214346],"study_design_scores_gemma":[0.0010707527,0.025805855,0.14416553,0.0007569731,0.00040249573,0.0035402216,0.03548845,0.4679765,0.26882353,0.00898229,0.04234806,0.0006393043],"about_ca_topic_score_codex":0.0006805226,"about_ca_topic_score_gemma":0.0009937542,"teacher_disagreement_score":0.011193395,"about_ca_system_score_codex":0.00067152,"about_ca_system_score_gemma":0.00079348637,"threshold_uncertainty_score":0.05919701},"labels":[],"label_agreement":null},{"id":"W4408295990","doi":"10.1007/978-3-031-83790-6_22","title":"Software Defects Prediction Using Generative Adversarial Network Based Data Balancing","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Adversarial system; Generative grammar; Generative adversarial network; Software; Data mining; Artificial intelligence; Programming language; Deep learning","score_opus":0.0589607147818736,"score_gpt":0.3099992995879466,"score_spread":0.251038584806073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408295990","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043508872,0.00042482262,0.9515895,0.00038011972,0.00013362501,0.00005753677,0.00014920559,0.001538313,0.0022180616],"genre_scores_gemma":[0.8968735,0.00018179265,0.09488359,0.00034781027,0.00014599288,0.00010977785,0.00071889895,0.00023412151,0.0065045888],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992269,0.00015134233,0.000036006928,0.00024634256,0.00024680604,0.00009257619],"domain_scores_gemma":[0.9976947,0.001261722,0.000198094,0.00032414394,0.00045479223,0.00006647067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015721807,0.00084370724,0.001025231,0.0010044699,0.0004035006,0.0008404688,0.0021620854,0.0013078095,0.001889821],"category_scores_gemma":[0.0044538076,0.00034775102,0.00071304373,0.00080807164,0.00066737266,0.001509384,0.0017233954,0.0015585768,0.0007160495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021501293,0.000096734046,0.0016079516,0.00003275866,0.000049700382,0.00008744128,0.00003998523,0.8099919,0.0042109755,0.004823133,0.0030558493,0.17578857],"study_design_scores_gemma":[0.0000022923887,0.000009836548,0.00008801097,0.0000017138898,0.0000033849658,0.00001027831,0.000002392806,0.9974911,0.00054909766,0.0017304114,0.000109711764,0.0000019082117],"about_ca_topic_score_codex":0.002368385,"about_ca_topic_score_gemma":0.0024101117,"teacher_disagreement_score":0.002368385,"about_ca_system_score_codex":0.0007686349,"about_ca_system_score_gemma":0.00045984826,"threshold_uncertainty_score":0.00831455},"labels":[],"label_agreement":null},{"id":"W4408693858","doi":"10.3390/software4020007","title":"Empirical Analysis of Data Sampling-Based Decision Forest Classifiers for Software Defect Prediction","year":2025,"lang":"en","type":"article","venue":"Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Decision tree; Computer science; Data mining; Sampling (signal processing); Machine learning; Software; Artificial intelligence; Statistics; Mathematics","score_opus":0.08940291661075161,"score_gpt":0.38240660768662454,"score_spread":0.29300369107587293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408693858","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.683655,0.0055844393,0.3044541,0.00079928833,0.00021768942,0.00032461158,0.002088605,0.00095622346,0.0019199898],"genre_scores_gemma":[0.9611551,0.0005421013,0.03465826,0.00011207479,0.00008814521,0.0001284564,0.0028380437,0.000033711247,0.00044406092],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961994,0.0015069456,0.00031193148,0.0006770787,0.00094673986,0.00035777924],"domain_scores_gemma":[0.96389526,0.028369324,0.0015319043,0.0016641166,0.0040526926,0.00048664003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014670443,0.0014787709,0.0016434095,0.0024492177,0.00073542877,0.0010825518,0.0015961614,0.0013555489,0.0007734395],"category_scores_gemma":[0.02797598,0.00031344665,0.0012097538,0.0016933437,0.0007803046,0.0018370557,0.0007474194,0.0018584878,0.00034821813],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011494478,0.0005018741,0.110483885,0.00037346,0.00027287644,0.00028459498,0.00018231657,0.69433194,0.0015913123,0.0036441518,0.0055380566,0.18164608],"study_design_scores_gemma":[0.000013385788,0.00008150097,0.005047658,0.000026272439,0.00002607286,0.000057386915,0.00003383804,0.9921862,0.00059135834,0.0015751699,0.00035026515,0.000010831421],"about_ca_topic_score_codex":0.007343754,"about_ca_topic_score_gemma":0.0057666404,"teacher_disagreement_score":0.014670443,"about_ca_system_score_codex":0.0011765087,"about_ca_system_score_gemma":0.0012605241,"threshold_uncertainty_score":0.07758564},"labels":[],"label_agreement":null},{"id":"W4408694007","doi":"10.1007/978-3-662-70810-1_3","title":"Anti-patterns and Code Smells for Multi-language Systems","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Polytechnique Montréal","funders":"","keywords":"Computer science; Programming language; Code smell; Code (set theory); Software quality; Software development; Software","score_opus":0.025156689610194335,"score_gpt":0.2929987955546165,"score_spread":0.26784210594442215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408694007","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2696167,0.0015381083,0.7014568,0.0019039367,0.00025476125,0.000055325876,0.00027937713,0.0051862947,0.019708697],"genre_scores_gemma":[0.8731691,0.0004839599,0.1120306,0.00026982898,0.000118963275,0.000054630054,0.00031408868,0.0011734033,0.012385357],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978915,0.0004018462,0.00020513395,0.0003510548,0.00095223583,0.00019824819],"domain_scores_gemma":[0.9863783,0.007855829,0.0019315586,0.0022855299,0.0012306819,0.00031811476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012385625,0.00043403413,0.00056166743,0.0011398188,0.00061813195,0.0018212107,0.0008367016,0.0012224086,0.0030463329],"category_scores_gemma":[0.013202045,0.0006899899,0.000677941,0.0010268516,0.0013867697,0.004622424,0.0017189983,0.0019087818,0.00058119936],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005332611,0.0002776455,0.02676435,0.0007832045,0.00011209935,0.00304655,0.0027830917,0.027763896,0.06366561,0.3540126,0.010483041,0.5097747],"study_design_scores_gemma":[0.000042181753,0.00027172235,0.007745124,0.00012899298,0.00006999863,0.00429896,0.00054497755,0.2987826,0.03658263,0.6327198,0.018725554,0.000087479595],"about_ca_topic_score_codex":0.0005301741,"about_ca_topic_score_gemma":0.00087715266,"teacher_disagreement_score":0.0030463329,"about_ca_system_score_codex":0.00048184028,"about_ca_system_score_gemma":0.0004661888,"threshold_uncertainty_score":0.010191023},"labels":[],"label_agreement":null},{"id":"W4408815477","doi":"10.1007/978-3-031-73143-3_4","title":"Natural Language Processing for Requirements Traceability","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Traceability; Computer science; Requirements traceability; Natural (archaeology); Programming language; Software engineering; Requirements analysis; History; Archaeology","score_opus":0.0268570263521775,"score_gpt":0.315352532398124,"score_spread":0.2884955060459465,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408815477","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017996398,0.0027913402,0.8896575,0.0013355026,0.00038769314,0.00013971124,0.00065432955,0.0035452812,0.09968903],"genre_scores_gemma":[0.04418114,0.005417493,0.7907374,0.00085807487,0.0003978569,0.00031813525,0.0037737438,0.0018126464,0.1525035],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993936,0.00013742498,0.000048346035,0.000101179896,0.00029064025,0.000028788401],"domain_scores_gemma":[0.99876213,0.0008474536,0.000038074668,0.00017093086,0.00016760525,0.000013772842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006863546,0.0007937121,0.00046361206,0.0014014248,0.00055147876,0.0019529641,0.0013256737,0.0007191498,0.028547965],"category_scores_gemma":[0.0026036273,0.00052180164,0.0010839542,0.0014443486,0.0011312654,0.0042855316,0.00086601946,0.0019204052,0.010472392],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003302625,0.00006541505,0.00007907949,0.0006084702,0.000014714127,0.0001647784,0.00041235398,0.0049235327,0.0067571574,0.34043714,0.06687717,0.57962716],"study_design_scores_gemma":[0.000013910052,0.000031035524,0.00023763029,0.00031731618,0.000022546785,0.0004372528,0.00019843546,0.04407283,0.009064761,0.59260404,0.3529654,0.00003477478],"about_ca_topic_score_codex":0.0025506455,"about_ca_topic_score_gemma":0.0029450157,"teacher_disagreement_score":0.028547965,"about_ca_system_score_codex":0.001128717,"about_ca_system_score_gemma":0.0009597995,"threshold_uncertainty_score":0.095502496},"labels":[],"label_agreement":null},{"id":"W4408840153","doi":"10.1007/978-3-031-73143-3_1","title":"Handbook on Natural Language Processing for Requirements Engineering: Overview","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Natural (archaeology); Software engineering; History; Archaeology","score_opus":0.03610718233777591,"score_gpt":0.3094514654281686,"score_spread":0.27334428309039266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408840153","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001118702,0.1265847,0.52209955,0.0029304742,0.0025770508,0.000585209,0.0071301246,0.013351789,0.32362235],"genre_scores_gemma":[0.0071915113,0.12358051,0.54377234,0.0021796199,0.0015518081,0.0008615168,0.017255463,0.0048415232,0.29876572],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992768,0.00008182525,0.000076561904,0.00011897684,0.0004146582,0.00003114748],"domain_scores_gemma":[0.99825484,0.0010284826,0.00005574036,0.00017182573,0.00044657334,0.000042487365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008447659,0.0014804102,0.0012948322,0.0042906743,0.0005799391,0.0035894816,0.0018533062,0.001136622,0.06739727],"category_scores_gemma":[0.0029371518,0.0011275869,0.00116375,0.0070295036,0.00072286493,0.0041211867,0.0011919525,0.0022694631,0.058166903],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017495313,0.00005381586,0.000055204724,0.0018463602,0.000017823659,0.00007947577,0.00019294638,0.0012485082,0.0032784464,0.033318717,0.27979717,0.680094],"study_design_scores_gemma":[0.00000649158,0.000012802168,0.0001748707,0.00043906496,0.000011682237,0.00031735614,0.000045032266,0.0010205398,0.000980156,0.026462104,0.97051144,0.00001848533],"about_ca_topic_score_codex":0.0021461963,"about_ca_topic_score_gemma":0.004098329,"teacher_disagreement_score":0.06739727,"about_ca_system_score_codex":0.0010598005,"about_ca_system_score_gemma":0.0018911067,"threshold_uncertainty_score":0.22546631},"labels":[],"label_agreement":null},{"id":"W4408844959","doi":"10.1145/3723876","title":"Automatically Improving LLM-based Verilog Generation using EDA Tool Feedback","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Design Automation of Electronic Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Verilog; Computer architecture; Embedded system; Parallel computing; Field-programmable gate array","score_opus":0.03001254172624539,"score_gpt":0.2764101130795037,"score_spread":0.2463975713532583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408844959","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07267676,0.00049999636,0.7940378,0.00047076162,0.0002503981,0.0003566451,0.0012392187,0.12539399,0.005074264],"genre_scores_gemma":[0.45402566,0.0002452773,0.52631646,0.00036558183,0.000040812396,0.00042289196,0.004029249,0.01080023,0.0037538074],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964275,0.0013035017,0.0002110654,0.00050121645,0.0013446742,0.00021213431],"domain_scores_gemma":[0.9863219,0.008923363,0.00058941956,0.0022297238,0.0016970596,0.0002385307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042972635,0.001848272,0.0007090073,0.0013982647,0.00039503627,0.0012843553,0.002441651,0.0008509111,0.0082140425],"category_scores_gemma":[0.025392218,0.0009376607,0.00088205317,0.00038997454,0.000744452,0.0018621455,0.002148386,0.0013899714,0.002503968],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011010801,0.00074492756,0.007613587,0.0015919879,0.0001375042,0.0010187012,0.0012922586,0.2928575,0.08762082,0.009319791,0.024501208,0.5722007],"study_design_scores_gemma":[0.00016696326,0.00023169041,0.00048603406,0.000112991,0.00004962473,0.00022533133,0.00013524876,0.9221196,0.057569694,0.005340028,0.013507612,0.00005514761],"about_ca_topic_score_codex":0.0018135072,"about_ca_topic_score_gemma":0.003074389,"teacher_disagreement_score":0.0082140425,"about_ca_system_score_codex":0.000994117,"about_ca_system_score_gemma":0.002194519,"threshold_uncertainty_score":0.027478755},"labels":[],"label_agreement":null},{"id":"W4409013497","doi":"10.1007/978-3-031-85628-0_22","title":"XAI-Based Assessment of Software Vulnerability Contributing Factors in Transformer Models","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Vulnerability (computing); Computer science; Software; Vulnerability assessment; Transformer; Engineering; Computer security; Medicine; Electrical engineering; Operating system","score_opus":0.048527713681525905,"score_gpt":0.3397813336764742,"score_spread":0.29125361999494825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409013497","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.101948015,0.0002790799,0.8798088,0.00011630519,0.0000318524,0.000100663274,0.00055471057,0.0019436921,0.015216864],"genre_scores_gemma":[0.85867655,0.0002497453,0.13616464,0.000019201256,0.000014393205,0.00010154433,0.0008257523,0.0002231289,0.0037250484],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99922025,0.00023040555,0.000038395796,0.00008510652,0.00036169187,0.00006424055],"domain_scores_gemma":[0.99733704,0.0015055819,0.00025290478,0.00031388327,0.0005222533,0.00006833696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017961765,0.0008705399,0.00059518433,0.0019842044,0.00033243152,0.0017481777,0.0010774554,0.000545054,0.004566198],"category_scores_gemma":[0.005722135,0.00029901564,0.0007458858,0.0011978899,0.0004719815,0.0020405126,0.0009211488,0.0007341086,0.0008060752],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021533723,0.00010680094,0.012688449,0.00019361019,0.000107778185,0.00020078663,0.00019867804,0.77710617,0.012425365,0.048213217,0.0023667938,0.14617702],"study_design_scores_gemma":[0.0000028699317,0.00003549855,0.0010413468,0.000013748452,0.000024878713,0.00006356196,0.00003239521,0.98775476,0.0023765122,0.008042271,0.00060409255,0.000007998099],"about_ca_topic_score_codex":0.0022905227,"about_ca_topic_score_gemma":0.0022909574,"teacher_disagreement_score":0.004566198,"about_ca_system_score_codex":0.0009093107,"about_ca_system_score_gemma":0.00069695764,"threshold_uncertainty_score":0.015275419},"labels":[],"label_agreement":null},{"id":"W4409014381","doi":"10.1109/access.2025.3556313","title":"Rethinking Technological Investment and Cost-Benefit: A Software Requirements Dependency Extraction Case Study","year":2025,"lang":"en","type":"article","venue":"IEEE Access","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Dependency (UML); Software; Investment (military); Risk analysis (engineering); Software engineering; Business; Programming language","score_opus":0.08223015109803639,"score_gpt":0.3846030789646008,"score_spread":0.3023729278665644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409014381","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9179696,0.0009752597,0.06598149,0.003209637,0.00005703874,0.00046682663,0.001712065,0.0005097435,0.009118253],"genre_scores_gemma":[0.89634466,0.0003436849,0.099879235,0.00028874783,0.00003075131,0.00020214928,0.0013228876,0.00013926481,0.0014486058],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.99115306,0.0054831775,0.00042607472,0.0006052772,0.0019729924,0.00035953298],"domain_scores_gemma":[0.9290494,0.062026214,0.001901423,0.0032866222,0.003086966,0.00064940465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009942599,0.0008241931,0.0005434096,0.0026677828,0.0010464366,0.0024043988,0.0014809804,0.0018481717,0.002183216],"category_scores_gemma":[0.037011597,0.00040406236,0.0010178345,0.0030982383,0.001293357,0.0033953954,0.0017265091,0.00249702,0.00052046817],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002271806,0.0042774505,0.13129295,0.002279258,0.00048468012,0.0066722785,0.0044810264,0.30622557,0.01679097,0.053316493,0.016140383,0.45576715],"study_design_scores_gemma":[0.00044322913,0.0020813623,0.047931425,0.00059172994,0.00032676465,0.002955663,0.0055870493,0.8207961,0.046132576,0.03664245,0.03627773,0.00023395795],"about_ca_topic_score_codex":0.0053076465,"about_ca_topic_score_gemma":0.012470424,"teacher_disagreement_score":0.009942599,"about_ca_system_score_codex":0.0031561824,"about_ca_system_score_gemma":0.001747226,"threshold_uncertainty_score":0.052582145},"labels":[],"label_agreement":null},{"id":"W4409405484","doi":"10.1007/s10664-025-10651-z","title":"Predicting the understandability of computational notebooks through code metrics analysis","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute","funders":"","keywords":"Computer science; Programming language; Code (set theory)","score_opus":0.03729922534070134,"score_gpt":0.3189262083606235,"score_spread":0.2816269830199222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409405484","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98981047,0.00011607091,0.0082276305,0.000078207595,0.000006786249,0.000028223394,0.0005652657,0.0003418314,0.00082551595],"genre_scores_gemma":[0.98748463,0.000083907325,0.010010483,0.000013829696,0.0000061965475,0.000026549107,0.0014907521,0.00009305151,0.00079062895],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9987047,0.0003083187,0.00010219162,0.0002254568,0.000583606,0.00007572612],"domain_scores_gemma":[0.95273274,0.030813415,0.0069952193,0.0026177894,0.006072269,0.0007686264],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0013753106,0.0005328371,0.00023920949,0.0036603597,0.0002773339,0.0013034137,0.00045914063,0.0006452331,0.0011697988],"category_scores_gemma":[0.04871589,0.00020501569,0.00036082647,0.0023303186,0.00034820303,0.0022713826,0.0006280015,0.00068229187,0.00050967094],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005330917,0.0004741016,0.6691598,0.00029050373,0.00019758783,0.0003175203,0.001443113,0.0476392,0.01891554,0.0022339418,0.0029730476,0.25582254],"study_design_scores_gemma":[0.00003951288,0.000847526,0.507492,0.0000915329,0.00010853278,0.00030217308,0.0009135983,0.45934057,0.022479469,0.0048443666,0.0034624878,0.00007831971],"about_ca_topic_score_codex":0.008177218,"about_ca_topic_score_gemma":0.011721239,"teacher_disagreement_score":0.9986247,"about_ca_system_score_codex":0.0007374247,"about_ca_system_score_gemma":0.0006207707,"threshold_uncertainty_score":0.016259253},"labels":[],"label_agreement":null},{"id":"W4409435353","doi":"10.1017/dsj.2025.7","title":"AI-driven FMEA: integration of large language models for faster and more accurate risk analysis","year":2025,"lang":"en","type":"article","venue":"Design Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Rimouski","funders":"Lunds Universitet","keywords":"Computer science; Risk analysis (engineering); Linguistics; Business; Philosophy","score_opus":0.031172969705433918,"score_gpt":0.3418963231019013,"score_spread":0.3107233533964674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409435353","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007244287,0.000078630845,0.9886232,0.00021943182,0.000022941049,0.000060580227,0.000069089845,0.0025759786,0.0011058898],"genre_scores_gemma":[0.25507346,0.00010755305,0.74247974,0.0001968153,0.000030210018,0.00016835003,0.00030174575,0.0005793675,0.0010626967],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978046,0.00110779,0.000120748846,0.00023236359,0.0006425053,0.00009197738],"domain_scores_gemma":[0.9910746,0.006487444,0.0005379955,0.0009201187,0.00083653856,0.00014327027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039284434,0.0009152234,0.00048456606,0.0011673733,0.0003541646,0.0018236977,0.0015877101,0.0008953222,0.0036156748],"category_scores_gemma":[0.014086688,0.0004451991,0.0011272739,0.00052784535,0.00088241335,0.002195956,0.0017777467,0.0018997579,0.00080183026],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021647714,0.00025074492,0.0022210553,0.00030369055,0.000145035,0.00025069565,0.0006959927,0.7280832,0.016495697,0.048498,0.0045339977,0.19830549],"study_design_scores_gemma":[0.000014118929,0.0000253791,0.00011019964,0.000022104443,0.000012661438,0.000029553272,0.000025410172,0.9826481,0.0024777888,0.012251723,0.0023707845,0.000012243504],"about_ca_topic_score_codex":0.0050853663,"about_ca_topic_score_gemma":0.004814717,"teacher_disagreement_score":0.0050853663,"about_ca_system_score_codex":0.0010898878,"about_ca_system_score_gemma":0.001745276,"threshold_uncertainty_score":0.020775795},"labels":[],"label_agreement":null},{"id":"W4409441986","doi":"10.1007/s10664-025-10655-9","title":"Predicting long time contributors with knowledge units of programming languages: an empirical study","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Empirical research; Programming language; Natural language processing; Data science; Artificial intelligence; Statistics; Mathematics","score_opus":0.017723677763842167,"score_gpt":0.3199314896089665,"score_spread":0.30220781184512435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409441986","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99918467,0.00004571381,0.0003041425,0.00006830019,0.0000046477226,0.000009010585,0.00008072743,0.000007384114,0.00029541907],"genre_scores_gemma":[0.9975005,0.00007342078,0.0004092199,0.000026880867,0.000015088724,0.000018636656,0.0002543497,0.000014731578,0.0016872106],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99784946,0.00094497314,0.00014663057,0.0003735636,0.00046374946,0.00022157791],"domain_scores_gemma":[0.8323547,0.11714985,0.022889588,0.0074130306,0.0069166566,0.013276225],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.003626236,0.0003620716,0.0003419942,0.0021263834,0.0010867858,0.0017946647,0.0013780391,0.0015549969,0.0067471876],"category_scores_gemma":[0.057838365,0.0005234166,0.00040958446,0.0014814213,0.0006475766,0.003115134,0.0016617001,0.0022511124,0.0016359518],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023950906,0.0011620885,0.9873892,0.000017224385,0.000042657153,0.00015280602,0.0020720745,0.00029263351,0.00027127433,0.0002207273,0.00035477863,0.0077849953],"study_design_scores_gemma":[0.00004918973,0.00058645144,0.9766533,0.00004047132,0.000109683824,0.00044051927,0.0063989027,0.0122944545,0.0008750968,0.00077310845,0.0017409769,0.000037956175],"about_ca_topic_score_codex":0.008350407,"about_ca_topic_score_gemma":0.010488621,"teacher_disagreement_score":0.9978736,"about_ca_system_score_codex":0.00054392515,"about_ca_system_score_gemma":0.001004037,"threshold_uncertainty_score":0.022571564},"labels":[],"label_agreement":null},{"id":"W4409602824","doi":"10.1007/s10115-025-02408-3","title":"Convolution neural network-AlexNet with gazelle optimization algorithm-based software defect prediction","year":2025,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Artificial neural network; Software; Convolution (computer science); Convolutional neural network; Algorithm; Artificial intelligence; Pattern recognition (psychology); Data mining; Machine learning","score_opus":0.006693307644077392,"score_gpt":0.21677554290158868,"score_spread":0.2100822352575113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409602824","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30647585,0.001892792,0.66660196,0.00074899045,0.00040523085,0.00012916191,0.0016675317,0.014669637,0.0074088597],"genre_scores_gemma":[0.9058656,0.000408354,0.08226555,0.00020391365,0.00007206571,0.000094864736,0.0027425922,0.00019123423,0.008155865],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99962866,0.000028641965,0.00001977341,0.000135748,0.000107961365,0.00007920995],"domain_scores_gemma":[0.9995906,0.00008033214,0.000042958523,0.00005483596,0.00020286163,0.000028446026],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005170647,0.0010035048,0.0010175756,0.0011278678,0.00036063849,0.0006342893,0.0016804906,0.0011025586,0.002330731],"category_scores_gemma":[0.0010649464,0.00036514853,0.000779052,0.0010072341,0.00031412815,0.0011958344,0.0007744034,0.0009932304,0.0009317607],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004506439,0.0004781147,0.009612527,0.00016232769,0.00020029467,0.00023095036,0.00005768349,0.35885626,0.014056915,0.0027566166,0.018866494,0.5942712],"study_design_scores_gemma":[0.0000050064214,0.000021064458,0.00059049705,0.000004652705,0.000012232809,0.00001618633,0.0000037134273,0.99674004,0.0019579355,0.00042059828,0.00022386218,0.0000042395486],"about_ca_topic_score_codex":0.023861652,"about_ca_topic_score_gemma":0.02133195,"teacher_disagreement_score":0.023861652,"about_ca_system_score_codex":0.00088068633,"about_ca_system_score_gemma":0.0013871111,"threshold_uncertainty_score":0.047445536},"labels":[],"label_agreement":null},{"id":"W4409608564","doi":"10.1007/s10664-025-10633-1","title":"Impact of extensions on browser performance: An empirical study on google chrome","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Empirical research; World Wide Web; Computer science; Business; Mathematics","score_opus":0.03426352058869829,"score_gpt":0.36338248516094646,"score_spread":0.32911896457224815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409608564","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9992987,0.000035169254,0.000026844786,0.000013348945,0.000001098722,0.00000360927,0.00005668859,0.000017838418,0.00054681377],"genre_scores_gemma":[0.9991542,0.000036929483,0.00013230633,0.000012455269,0.000004441309,0.0000033687695,0.00022249398,0.000015526546,0.0004182949],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99849534,0.00041932787,0.00008441467,0.00020810703,0.0005435319,0.00024934552],"domain_scores_gemma":[0.95974964,0.026475245,0.0041555665,0.0023864415,0.0046355603,0.0025975737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017919289,0.00040309326,0.00029155792,0.0012650386,0.000590431,0.0017218293,0.0005120182,0.0007405186,0.0019710672],"category_scores_gemma":[0.02233327,0.00020182997,0.00033267087,0.0018180453,0.0006697505,0.0025566772,0.0005102329,0.0013983413,0.00057736103],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017317514,0.0046101515,0.9472607,0.00017554412,0.00018900499,0.00058383215,0.0024797844,0.0035468396,0.0061716354,0.0006561125,0.002189572,0.030405173],"study_design_scores_gemma":[0.000048995433,0.0011674119,0.9849744,0.000033252913,0.00013377774,0.00027445614,0.0021659639,0.007830023,0.0019470769,0.00018691472,0.0011960266,0.000041679075],"about_ca_topic_score_codex":0.01303917,"about_ca_topic_score_gemma":0.013376428,"teacher_disagreement_score":0.01303917,"about_ca_system_score_codex":0.00077142444,"about_ca_system_score_gemma":0.00082823506,"threshold_uncertainty_score":0.02592653},"labels":[],"label_agreement":null},{"id":"W4409642872","doi":"10.5220/0013211700003928","title":"Recommender Systems Approaches for Software Defect Prediction: A Comparative Study","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Computer science; Recommender system; Software bug; Software; Software engineering; Machine learning; Programming language","score_opus":0.12016000022530869,"score_gpt":0.3276400159924681,"score_spread":0.2074800157671594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409642872","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77634114,0.07660898,0.112821765,0.004327418,0.00044613006,0.0004787564,0.0009241685,0.0008526616,0.027199024],"genre_scores_gemma":[0.93739927,0.011691114,0.04668919,0.0002554643,0.00027716983,0.00008405813,0.00059942383,0.000045556615,0.0029586894],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9908852,0.0052443137,0.00038173667,0.00078963063,0.0024996458,0.00019954331],"domain_scores_gemma":[0.8784586,0.10874432,0.0019166749,0.0029614586,0.0071380446,0.0007808214],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015067235,0.0010766597,0.001570588,0.0059152003,0.00094837055,0.0028550932,0.0022323923,0.002296243,0.0023711042],"category_scores_gemma":[0.040987957,0.0004322211,0.001272736,0.0054934695,0.00067115086,0.0046016895,0.0010651047,0.0015964631,0.0008069609],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023662362,0.0026899616,0.10549381,0.002083416,0.0024445462,0.00014277459,0.0016254393,0.04509437,0.00087173085,0.008287846,0.0047307713,0.8241691],"study_design_scores_gemma":[0.00060724263,0.0073229987,0.16746552,0.0010636614,0.0027999065,0.0007266648,0.0034171003,0.783819,0.002315354,0.014685515,0.015415914,0.00036107455],"about_ca_topic_score_codex":0.020621037,"about_ca_topic_score_gemma":0.018832076,"teacher_disagreement_score":0.020621037,"about_ca_system_score_codex":0.0019480101,"about_ca_system_score_gemma":0.0012074261,"threshold_uncertainty_score":0.07968408},"labels":[],"label_agreement":null},{"id":"W4409814328","doi":"10.1016/j.procs.2025.03.107","title":"Automated UML Visualization of Software Ecosystems: Tracking Versions, Dependencies, and Security Updates","year":2025,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Unified Modeling Language; Visualization; Software; Software engineering; UML tool; Programming language; Data mining","score_opus":0.010078912977670608,"score_gpt":0.2758635128138951,"score_spread":0.2657845998362245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409814328","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12292284,0.0005627074,0.83792776,0.0010275805,0.00013166858,0.00014279915,0.0015317709,0.02905103,0.0067017875],"genre_scores_gemma":[0.48561388,0.0006293231,0.5075089,0.00012849984,0.000039382296,0.0001304341,0.0017435561,0.0020494144,0.0021566025],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99895704,0.00038561804,0.000095178424,0.00013854692,0.0003694553,0.00005420988],"domain_scores_gemma":[0.9931305,0.0029724576,0.0011154591,0.0012520967,0.0012386965,0.00029071403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025829438,0.00080232596,0.00034017095,0.004008054,0.0005817519,0.0029200658,0.0007482241,0.0008702559,0.0022385132],"category_scores_gemma":[0.0117670195,0.0005716932,0.0005199515,0.0015007686,0.00038593536,0.0026262226,0.0018265139,0.0010164928,0.0006583049],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006884039,0.0004039472,0.0665158,0.00077452126,0.00019337492,0.001556759,0.017819345,0.073122516,0.0719284,0.05295503,0.026991894,0.68705],"study_design_scores_gemma":[0.000080782636,0.0001606075,0.017064366,0.0003848668,0.00014564731,0.0008782443,0.0018830117,0.8308349,0.047231898,0.026475787,0.07466936,0.00019055376],"about_ca_topic_score_codex":0.0048530307,"about_ca_topic_score_gemma":0.0077385325,"teacher_disagreement_score":0.0048530307,"about_ca_system_score_codex":0.0007959252,"about_ca_system_score_gemma":0.001522494,"threshold_uncertainty_score":0.013660014},"labels":[],"label_agreement":null},{"id":"W4410552808","doi":"10.1109/saner64311.2025.00062","title":"On the Performance of Large Language Models for Code Change Intent Classification","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Concordia University","funders":"","keywords":"Computer science; Programming language; Code (set theory); Natural language processing; Artificial intelligence","score_opus":0.06768386307145074,"score_gpt":0.3200188008704129,"score_spread":0.25233493779896216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410552808","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6775732,0.009071436,0.28875476,0.0033514244,0.00065759715,0.00038980873,0.0025193545,0.011483435,0.0061989473],"genre_scores_gemma":[0.9261355,0.00066829537,0.0670052,0.00047769272,0.00020116112,0.00016111947,0.0033745498,0.00025238693,0.0017241925],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99342465,0.0033690822,0.00044070944,0.0014037851,0.0010271854,0.00033458637],"domain_scores_gemma":[0.9538935,0.038502567,0.0018072521,0.0020491304,0.0031030406,0.000644591],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0101242885,0.002396447,0.0014120021,0.004514462,0.0009615696,0.0021224115,0.0016374079,0.001746767,0.0013407985],"category_scores_gemma":[0.03763637,0.0005303816,0.0016430081,0.002129982,0.0006950238,0.0033935823,0.001522228,0.0027924958,0.0014836673],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015349535,0.0010142858,0.048621148,0.00046935742,0.0004996267,0.00028002594,0.00048230158,0.38874987,0.0042906185,0.003054671,0.012386157,0.538617],"study_design_scores_gemma":[0.000014935364,0.00010414574,0.0013988637,0.000016454678,0.000025841064,0.000030874304,0.000047382084,0.99596417,0.00085993897,0.0011817736,0.0003404391,0.000015164139],"about_ca_topic_score_codex":0.021240775,"about_ca_topic_score_gemma":0.019369336,"teacher_disagreement_score":0.021240775,"about_ca_system_score_codex":0.0019995025,"about_ca_system_score_gemma":0.0018572755,"threshold_uncertainty_score":0.05354297},"labels":[],"label_agreement":null},{"id":"W4410553379","doi":"10.1109/saner64311.2025.00025","title":"Analysing Software Supply Chains of Infrastructure as Code: Extraction of Ansible Plugin Dependencies","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Vlaamse regering; Fonds Wetenschappelijk Onderzoek","keywords":"Plug-in; Computer science; Code (set theory); Software; Supply chain; Extraction (chemistry); Database; Software engineering; Programming language; Business; Chemistry","score_opus":0.009035467800809521,"score_gpt":0.2888955027134334,"score_spread":0.2798600349126239,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410553379","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40626663,0.00087501353,0.552621,0.0004720112,0.00007488106,0.0007584323,0.0031768205,0.026774041,0.008981287],"genre_scores_gemma":[0.46535584,0.0008172964,0.51097524,0.00016418647,0.000046816258,0.0006947714,0.0130691705,0.0032236176,0.005652981],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99255884,0.0011086924,0.0007690756,0.0011973141,0.004032724,0.0003332936],"domain_scores_gemma":[0.96748394,0.014419302,0.0057949657,0.0052922713,0.00665181,0.0003577692],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036317413,0.0013169209,0.00062085467,0.0093473205,0.0012235822,0.0016498899,0.0009955604,0.00080626574,0.0012469954],"category_scores_gemma":[0.03208323,0.0011133639,0.0009171545,0.004016036,0.0007892363,0.0027625312,0.0022295255,0.0014261758,0.0009402372],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004290345,0.0004863054,0.1653791,0.0016587821,0.00019856742,0.0034312957,0.008245887,0.016693108,0.05744034,0.014942346,0.011160224,0.719935],"study_design_scores_gemma":[0.00005577177,0.0002747865,0.1448816,0.0010473141,0.00032702828,0.0037023865,0.003090732,0.50454414,0.19988333,0.019731265,0.1222197,0.0002419572],"about_ca_topic_score_codex":0.0070685893,"about_ca_topic_score_gemma":0.009994416,"teacher_disagreement_score":0.0093473205,"about_ca_system_score_codex":0.001138738,"about_ca_system_score_gemma":0.0038111326,"threshold_uncertainty_score":0.019206703},"labels":[],"label_agreement":null},{"id":"W4410705961","doi":"10.1049/sfw2/8832164","title":"Predicting Software Perfection Through Advanced Models to Uncover and Prevent Defects","year":2025,"lang":"en","type":"article","venue":"IET Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"University of Johannesburg","keywords":"Perfection; Computer science; Software engineering; Software; Programming language; Epistemology; Philosophy","score_opus":0.015748346247657422,"score_gpt":0.27230584308185884,"score_spread":0.2565574968342014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410705961","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.59292036,0.0016447837,0.39835784,0.0006299177,0.00009374389,0.00010641657,0.0008523045,0.0032422957,0.00215231],"genre_scores_gemma":[0.95330834,0.00033187648,0.04471686,0.0000653849,0.000026537271,0.000041848325,0.000731372,0.00003904869,0.00073863624],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993118,0.00019651029,0.000047059944,0.0001669327,0.00018551307,0.000092256894],"domain_scores_gemma":[0.99605227,0.00230498,0.0005675402,0.0002546384,0.0007353079,0.000085202446],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025525193,0.0012312434,0.0007600974,0.0021419607,0.00022138705,0.0008279217,0.0009177755,0.000809519,0.00071149203],"category_scores_gemma":[0.006694401,0.00031213157,0.000697016,0.0010035335,0.00030647762,0.0014601322,0.0005380638,0.0009080094,0.000542835],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001727544,0.0003689432,0.04965375,0.0001251582,0.00013159266,0.00010126875,0.00005758164,0.80514526,0.0025174539,0.0008380863,0.0021721253,0.13871603],"study_design_scores_gemma":[0.0000033328552,0.00005548691,0.003019186,0.000011831005,0.000012239364,0.00002217034,0.0000104550145,0.9954659,0.00064763,0.0005673688,0.00017889442,0.000005472832],"about_ca_topic_score_codex":0.0060514556,"about_ca_topic_score_gemma":0.007690209,"teacher_disagreement_score":0.0060514556,"about_ca_system_score_codex":0.00051707774,"about_ca_system_score_gemma":0.0008842435,"threshold_uncertainty_score":0.0134992},"labels":[],"label_agreement":null},{"id":"W4410771641","doi":"10.1145/3736405","title":"Improving Code Reviewer Recommendation: Accuracy, Latency, Workload, and Bystanders","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Workload; Code (set theory); Latency (audio); Operating system; Telecommunications; Programming language","score_opus":0.07592995547935187,"score_gpt":0.34187308048692294,"score_spread":0.26594312500757106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410771641","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.777763,0.020869741,0.1025967,0.012106114,0.0044838353,0.05863597,0.004866925,0.0042659203,0.014411741],"genre_scores_gemma":[0.8156888,0.0021441944,0.12221518,0.0035422407,0.0011993482,0.050485868,0.0011259139,0.00035243377,0.003246006],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8089734,0.13829252,0.02031814,0.009471051,0.020945083,0.0019998176],"domain_scores_gemma":[0.3070288,0.5809882,0.050146967,0.029349271,0.027425367,0.0050613466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.16866674,0.0014578751,0.0023266077,0.0018592115,0.0014210037,0.0031819784,0.002646645,0.0036766604,0.0059917555],"category_scores_gemma":[0.47928455,0.0011046118,0.0034530198,0.0017129154,0.0020354302,0.004657249,0.0019024414,0.0033973372,0.0023069794],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.1477088,0.024839044,0.053415064,0.025378067,0.008585053,0.00018486273,0.0025294742,0.008092523,0.009078204,0.0025954756,0.017614447,0.699979],"study_design_scores_gemma":[0.14496918,0.5530471,0.12469719,0.0075039677,0.019370629,0.00096452556,0.0018286635,0.04792901,0.041820057,0.015069014,0.041543074,0.0012575185],"about_ca_topic_score_codex":0.0015005564,"about_ca_topic_score_gemma":0.0021938924,"teacher_disagreement_score":0.16866674,"about_ca_system_score_codex":0.0022541347,"about_ca_system_score_gemma":0.005142894,"threshold_uncertainty_score":0.8920056},"labels":[],"label_agreement":null},{"id":"W4411016400","doi":"10.1145/3714393.3726509","title":"VulPatrol: Interprocedural Vulnerability Detection and Localization through Semantic Graph Learning","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Vulnerability (computing); Graph; Artificial intelligence; Natural language processing; Theoretical computer science; Computer security","score_opus":0.01164193984985394,"score_gpt":0.2713651300038017,"score_spread":0.25972319015394774,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411016400","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12878335,0.0022318505,0.7176939,0.0011446986,0.00022540217,0.0003346279,0.00646496,0.13646999,0.006651244],"genre_scores_gemma":[0.5809411,0.0010227165,0.3762551,0.00096701685,0.00008100548,0.0002945162,0.030855082,0.0026763838,0.0069071394],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99941885,0.00008306936,0.000024659617,0.00024950548,0.00016055618,0.00006336705],"domain_scores_gemma":[0.9991542,0.0003028545,0.00010356586,0.00025445936,0.00013789731,0.00004692579],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00056569633,0.0018496396,0.00062629,0.002055173,0.0004658345,0.00091840274,0.002579152,0.0013118671,0.0020952926],"category_scores_gemma":[0.0026940831,0.000534297,0.0013835765,0.0010667599,0.0007729598,0.0027913991,0.0014582146,0.0017581142,0.0010214658],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034856156,0.0005257616,0.009988494,0.00057108916,0.00033559927,0.00035191502,0.00016095764,0.32461467,0.015487024,0.008208805,0.04757155,0.5918355],"study_design_scores_gemma":[0.000020245425,0.000055424316,0.00077284523,0.000023382347,0.00002757172,0.00007081357,0.000024297266,0.97998536,0.005645121,0.009598133,0.003761087,0.000015809808],"about_ca_topic_score_codex":0.016203176,"about_ca_topic_score_gemma":0.028102241,"teacher_disagreement_score":0.016203176,"about_ca_system_score_codex":0.0016978363,"about_ca_system_score_gemma":0.0019036306,"threshold_uncertainty_score":0.03221768},"labels":[],"label_agreement":null},{"id":"W4411095886","doi":"10.1007/978-3-031-94590-8_19","title":"Studying Workarounds in Software Forms: An Experimental Protocol","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Workaround; Protocol (science); Computer science; Programming language; Medicine","score_opus":0.022313338324837905,"score_gpt":0.29989875253414916,"score_spread":0.27758541420931127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411095886","genre_codex":"protocol","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":"protocol","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35978305,0.00021032912,0.032074116,0.0005893363,0.0005707013,0.5700217,0.004537216,0.00071825203,0.0314954],"genre_scores_gemma":[0.0639605,0.00014377295,0.0220096,0.0002806122,0.00009065466,0.9055342,0.00095078145,0.000119771714,0.0069101406],"study_design_codex":"nonrandomized_trial","study_design_gemma":"not_applicable","domain_scores_codex":[0.9751066,0.008609877,0.004279527,0.0039352085,0.0056784884,0.0023902953],"domain_scores_gemma":[0.90041673,0.05503247,0.0069094244,0.018656319,0.015797185,0.0031879004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018966552,0.0028224317,0.0024344665,0.0023974834,0.0052430257,0.0028234893,0.0040747444,0.005502457,0.044424046],"category_scores_gemma":[0.063068286,0.0034947312,0.001769821,0.002803231,0.006189337,0.003141396,0.005320882,0.0069758575,0.010365862],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.21303582,0.32315043,0.016993647,0.009691197,0.00047896206,0.000944382,0.037447583,0.016483705,0.10408841,0.07205761,0.020454112,0.18517415],"study_design_scores_gemma":[0.20038426,0.29362327,0.08096396,0.0020681447,0.0023550196,0.00044716403,0.012828733,0.019409874,0.11796812,0.07764266,0.19073217,0.0015766901],"about_ca_topic_score_codex":0.0025223491,"about_ca_topic_score_gemma":0.0027693403,"teacher_disagreement_score":0.044424046,"about_ca_system_score_codex":0.00387828,"about_ca_system_score_gemma":0.007757023,"threshold_uncertainty_score":0.14861327},"labels":[],"label_agreement":null},{"id":"W4411232430","doi":"10.1109/tse.2025.3579574","title":"BLAZE: Cross-Language and Cross-Project Bug Localization via Dynamic Chunking and Hard Example Learning","year":2025,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Waterloo","funders":"","keywords":"Computer science; Chunking (psychology); Programming language; Software engineering; Artificial intelligence; Natural language processing","score_opus":0.010397706086851791,"score_gpt":0.2798341540061944,"score_spread":0.26943644791934257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411232430","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03354417,0.00037777776,0.82306075,0.00073211757,0.00015343718,0.0003617528,0.0008240105,0.13839372,0.0025522944],"genre_scores_gemma":[0.15863477,0.00014828444,0.82728803,0.00056757027,0.00003934141,0.00038218836,0.0030003944,0.0051984577,0.004740914],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99599755,0.0010708267,0.00027367828,0.0013777579,0.0010051976,0.00027503516],"domain_scores_gemma":[0.9898369,0.0035196808,0.00092071656,0.0040808944,0.0011763826,0.00046540567],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046475027,0.00252505,0.0013192251,0.0022569345,0.00084734894,0.0018533549,0.0056094993,0.0020890029,0.004482385],"category_scores_gemma":[0.020092715,0.0016703074,0.0016423232,0.0013409283,0.0015376498,0.007209259,0.008826314,0.0038028175,0.0025363814],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075177324,0.0007254828,0.013602375,0.00044307145,0.00031309595,0.00050828146,0.0015434494,0.08985984,0.020729594,0.008526217,0.040288653,0.82270813],"study_design_scores_gemma":[0.00015235863,0.00031106337,0.0020734502,0.00004698885,0.000067112545,0.0002484189,0.0002824485,0.9424685,0.02002176,0.019357998,0.014877562,0.000092362345],"about_ca_topic_score_codex":0.007857844,"about_ca_topic_score_gemma":0.017470857,"teacher_disagreement_score":0.007857844,"about_ca_system_score_codex":0.0012381421,"about_ca_system_score_gemma":0.0029410685,"threshold_uncertainty_score":0.02457869},"labels":[],"label_agreement":null},{"id":"W4411233811","doi":"10.1109/chase66643.2025.00029","title":"How Programmers Interact with Multimodal Software Documentation","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Documentation; Computer science; Software engineering; Software documentation; Programming language; Software; Software development; Software construction; Human–computer interaction","score_opus":0.008393029694554516,"score_gpt":0.2749351296675084,"score_spread":0.26654209997295386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411233811","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9320898,0.00040056935,0.050692815,0.0012217812,0.000031270138,0.00008988575,0.00003429947,0.0008054921,0.01463412],"genre_scores_gemma":[0.9674287,0.00037760113,0.026607463,0.00041959065,0.000030594925,0.00008859557,0.00008901886,0.00023082118,0.0047276937],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9920542,0.0059036715,0.00023370692,0.0004628165,0.0008883132,0.00045738506],"domain_scores_gemma":[0.9786451,0.016089976,0.0017461403,0.00090947544,0.0017127728,0.0008966212],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005366965,0.00072729343,0.00037941313,0.0014211595,0.0013041414,0.0046301335,0.00094719435,0.0020705105,0.004427102],"category_scores_gemma":[0.043234807,0.0005682817,0.00042756135,0.0005726339,0.0008805345,0.0051439432,0.0030643316,0.0012423445,0.0011115596],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005630029,0.000883525,0.07493563,0.00091176457,0.00013367037,0.0029277015,0.5364256,0.0025136352,0.061934084,0.005644657,0.0073495624,0.3057772],"study_design_scores_gemma":[0.00038885477,0.0033134602,0.12154299,0.0020597044,0.0007334485,0.014336634,0.538861,0.054800328,0.036707323,0.023003213,0.20358804,0.0006650351],"about_ca_topic_score_codex":0.0006177029,"about_ca_topic_score_gemma":0.000957537,"teacher_disagreement_score":0.005366965,"about_ca_system_score_codex":0.0004413764,"about_ca_system_score_gemma":0.0006465351,"threshold_uncertainty_score":0.028383553},"labels":[],"label_agreement":null},{"id":"W4411271062","doi":"10.1109/msr66628.2025.00066","title":"Characterizing Packages for Vulnerability Prediction","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Vulnerability (computing); Computer science; Vulnerability assessment; Computer security","score_opus":0.02003681168579495,"score_gpt":0.29637691353906054,"score_spread":0.2763401018532656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411271062","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82084954,0.0031120265,0.11036009,0.0011670021,0.0002335563,0.00038965474,0.042259984,0.012519294,0.009108863],"genre_scores_gemma":[0.8619935,0.00082354835,0.06732871,0.000244716,0.00011399653,0.00034761048,0.064007364,0.0009853506,0.004155226],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99874413,0.00016521872,0.00010271097,0.00035134854,0.00045030355,0.00018619176],"domain_scores_gemma":[0.9937856,0.0021960535,0.0014036269,0.0010590156,0.0012126209,0.0003431005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011678911,0.0013351446,0.0006139471,0.005070218,0.0005953707,0.0010470503,0.00091103296,0.0012849902,0.0014212461],"category_scores_gemma":[0.009292845,0.0003745406,0.0012969206,0.002830062,0.0005446409,0.002609975,0.0013952472,0.0013931898,0.0019427653],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060798746,0.0004920372,0.5968572,0.00068467733,0.0002231029,0.0015978536,0.0010175296,0.031679425,0.015499292,0.0044845063,0.077430345,0.26942602],"study_design_scores_gemma":[0.0000578142,0.0005784875,0.38100415,0.00028541192,0.0003301395,0.0045475103,0.0011699311,0.46251455,0.022705369,0.01934214,0.107290536,0.00017395588],"about_ca_topic_score_codex":0.005865778,"about_ca_topic_score_gemma":0.008369614,"teacher_disagreement_score":0.005865778,"about_ca_system_score_codex":0.00071270677,"about_ca_system_score_gemma":0.00082538533,"threshold_uncertainty_score":0.011663258},"labels":[],"label_agreement":null},{"id":"W4411271466","doi":"10.1109/msr66628.2025.00075","title":"CoUpJava: A Dataset of Code Upgrade Histories in Open-Source Java Repositories","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; University of Waterloo","keywords":"Computer science; Java; Upgrade; Open source; Code (set theory); Source code; Programming language; Operating system; Database; World Wide Web; Software; Set (abstract data type)","score_opus":0.01965083016863603,"score_gpt":0.30510215557404463,"score_spread":0.28545132540540863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411271466","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24768713,0.004812697,0.008547114,0.0009091213,0.0004068712,0.0002987958,0.7141512,0.016898766,0.0062882616],"genre_scores_gemma":[0.067614034,0.00047412733,0.010306974,0.00020902116,0.000054597716,0.0002502546,0.9188364,0.0007299982,0.0015245355],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99579597,0.00064113975,0.00048981607,0.0012873503,0.0014319951,0.00035373974],"domain_scores_gemma":[0.9877935,0.0035498044,0.0013967359,0.0036432317,0.0028977038,0.0007190517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019798563,0.0012646862,0.0007563236,0.007238237,0.000984905,0.0015899106,0.0021691436,0.0021674186,0.0015787737],"category_scores_gemma":[0.015582074,0.00060529466,0.0012275793,0.0066072973,0.0007319312,0.0020921254,0.002147548,0.0018755472,0.0029589378],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014586333,0.0008917621,0.20673992,0.004392099,0.0007514906,0.0020112735,0.0019304568,0.017126992,0.016026394,0.004197142,0.5679237,0.17655014],"study_design_scores_gemma":[0.00035479458,0.00044182656,0.37363407,0.00068850745,0.00030614153,0.0025781412,0.0010927661,0.04148353,0.01231427,0.00444667,0.5623113,0.00034802814],"about_ca_topic_score_codex":0.011993194,"about_ca_topic_score_gemma":0.0254814,"teacher_disagreement_score":0.011993194,"about_ca_system_score_codex":0.0010073872,"about_ca_system_score_gemma":0.0015284051,"threshold_uncertainty_score":0.023846805},"labels":[],"label_agreement":null},{"id":"W4411271537","doi":"10.1109/msr66628.2025.00087","title":"Investigating the Understandability of Review Comments on Code Change Requests","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Programming language; Code (set theory); Software engineering","score_opus":0.12091770400538404,"score_gpt":0.3694244875331656,"score_spread":0.24850678352778155,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411271537","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95492256,0.0026877872,0.02935939,0.001612946,0.00041257383,0.0012664379,0.0033236002,0.0013872928,0.0050274343],"genre_scores_gemma":[0.95668846,0.0009619869,0.032267425,0.00078018795,0.0003883055,0.0015899586,0.0043985946,0.0005006577,0.0024244308],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.8750958,0.055643152,0.018856289,0.01256144,0.03526321,0.002580084],"domain_scores_gemma":[0.15648498,0.6285764,0.1133188,0.01499179,0.08423312,0.0023948932],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07771897,0.0011330022,0.0011913952,0.0119481385,0.0019246103,0.0043628667,0.0012382882,0.0017818647,0.0015532892],"category_scores_gemma":[0.56152385,0.00061401824,0.0009795807,0.0056252005,0.0015272734,0.0055209096,0.0036567233,0.0018044731,0.001347472],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019075533,0.00038540523,0.61574703,0.0071342182,0.0007671049,0.0010362684,0.07435186,0.0020460864,0.019399945,0.0013191515,0.01975237,0.25615296],"study_design_scores_gemma":[0.00022348377,0.0015635695,0.83042663,0.004276235,0.0010665861,0.0026824588,0.032866728,0.035152238,0.025334204,0.003385533,0.062237658,0.0007846144],"about_ca_topic_score_codex":0.0035930597,"about_ca_topic_score_gemma":0.0068986714,"teacher_disagreement_score":0.07771897,"about_ca_system_score_codex":0.0026213038,"about_ca_system_score_gemma":0.0036880053,"threshold_uncertainty_score":0.41102213},"labels":[],"label_agreement":null},{"id":"W4411271539","doi":"10.1109/msr66628.2025.00119","title":"DPy: Code Smells Detection Tool for Python","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Python (programming language); Computer science; Programming language; Code smell; Operating system; Software quality; Software; Software development","score_opus":0.013989831227347007,"score_gpt":0.2868717867273277,"score_spread":0.27288195549998073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411271539","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017780337,0.00035199345,0.28599262,0.00052812666,0.00018496605,0.0008925206,0.012604538,0.67698413,0.0046806834],"genre_scores_gemma":[0.17715204,0.000604766,0.66839284,0.0010744918,0.00011819708,0.0030437591,0.041374665,0.092316255,0.015923029],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99574643,0.0005638247,0.00046184578,0.0006825475,0.0022945919,0.00025070805],"domain_scores_gemma":[0.9864058,0.006223164,0.002376657,0.0020848985,0.002472088,0.00043732784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00432173,0.0017846581,0.00082136045,0.0045466637,0.00082759245,0.0016956499,0.002595491,0.001224817,0.012505595],"category_scores_gemma":[0.022069067,0.0014611474,0.001294047,0.0017671773,0.0010918826,0.0040302337,0.004049979,0.0023825238,0.007627715],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011409858,0.00055184227,0.03240908,0.0031782251,0.00024280668,0.0012652794,0.0025928838,0.0061596055,0.03767155,0.0068670227,0.3986979,0.5092229],"study_design_scores_gemma":[0.0007191867,0.0012582387,0.058470096,0.0014864693,0.00026227933,0.004631241,0.00092524174,0.22347379,0.1877878,0.023602644,0.4963842,0.0009988156],"about_ca_topic_score_codex":0.0026070187,"about_ca_topic_score_gemma":0.0031676951,"teacher_disagreement_score":0.012505595,"about_ca_system_score_codex":0.0007293094,"about_ca_system_score_gemma":0.0026652392,"threshold_uncertainty_score":0.041835368},"labels":[],"label_agreement":null},{"id":"W4411271588","doi":"10.1109/msr66628.2025.00038","title":"Combining Large Language Models with Static Analyzers for Code Review Generation","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Programming language; Code generation; Code (set theory); Key (lock); Operating system","score_opus":0.030497533779869343,"score_gpt":0.324670604131976,"score_spread":0.29417307035210666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411271588","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039875403,0.0015086676,0.8613852,0.001192746,0.00030515142,0.00072053383,0.004005616,0.08852391,0.002482602],"genre_scores_gemma":[0.28620327,0.00057302695,0.69041497,0.0009771166,0.00023169756,0.0008865186,0.0153328655,0.0024709054,0.0029096673],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99438536,0.002249931,0.0005348115,0.0015050258,0.0011291418,0.0001958011],"domain_scores_gemma":[0.9692617,0.019262506,0.002059332,0.003703473,0.0051769284,0.0005360944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005362904,0.0017735097,0.0011350336,0.00486951,0.00062885473,0.002217684,0.0026800984,0.0013866661,0.0031140551],"category_scores_gemma":[0.037812706,0.00082414434,0.0015006422,0.0020852636,0.0006452654,0.0034918962,0.0020365547,0.0028181095,0.004962697],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040054682,0.00045268657,0.013846827,0.0007834264,0.0002616029,0.00045412168,0.0006157353,0.05740172,0.01743399,0.003561355,0.033879034,0.87090904],"study_design_scores_gemma":[0.000100418256,0.0001632334,0.0017852007,0.00008445129,0.00012277013,0.00020624521,0.00013296421,0.9578151,0.015858253,0.011261585,0.012408168,0.0000616317],"about_ca_topic_score_codex":0.0066567464,"about_ca_topic_score_gemma":0.0191902,"teacher_disagreement_score":0.0066567464,"about_ca_system_score_codex":0.0014880798,"about_ca_system_score_gemma":0.0037355516,"threshold_uncertainty_score":0.028362036},"labels":[],"label_agreement":null},{"id":"W4411271665","doi":"10.1109/msr66628.2025.00069","title":"It Works (only) on My Machine: A Study on Reproducibility Smells in Ansible Scripts","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Scripting language; Computer science; Reproducibility; Programming language; Mathematics; Statistics","score_opus":0.029881526054882708,"score_gpt":0.3231813417224383,"score_spread":0.2932998156675556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411271665","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99692816,0.00018015366,0.001062055,0.00020488257,0.0000085519305,0.000030088442,0.00007284214,0.000052030577,0.0014611316],"genre_scores_gemma":[0.9965779,0.00019307909,0.001718105,0.00015293724,0.000013882643,0.000060343893,0.00027445657,0.000107060834,0.000902192],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9911361,0.0042294227,0.0009284857,0.0009475871,0.0022586023,0.00049981254],"domain_scores_gemma":[0.8271181,0.110763796,0.033852395,0.009828706,0.014627421,0.0038095207],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012705607,0.00037208197,0.0003913012,0.002241383,0.0015424018,0.0029398773,0.0011175377,0.0007520399,0.0012792746],"category_scores_gemma":[0.08281388,0.00041858177,0.00048688825,0.0023264417,0.001907546,0.004762307,0.002015816,0.0013844155,0.00054143975],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005952899,0.0011326164,0.69002485,0.00092794414,0.00013955549,0.0023641395,0.18203366,0.001052132,0.006057764,0.00242783,0.0049583022,0.108285874],"study_design_scores_gemma":[0.000043982825,0.0015223516,0.8460306,0.00076077046,0.00018187102,0.0017197548,0.112690784,0.008577446,0.0051581524,0.0024635394,0.020664623,0.00018612643],"about_ca_topic_score_codex":0.0034055035,"about_ca_topic_score_gemma":0.005898283,"teacher_disagreement_score":0.9872944,"about_ca_system_score_codex":0.0011674373,"about_ca_system_score_gemma":0.0015469416,"threshold_uncertainty_score":0.06719446},"labels":[],"label_agreement":null},{"id":"W4411271765","doi":"10.1109/icse-companion66252.2025.00072","title":"On the Automation of Code Review Tasks Through Cross-Task Knowledge Distillation","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Automation; Task (project management); Distillation; Code (set theory); Programming language; Software engineering; Human–computer interaction; Engineering; Systems engineering; Chromatography; Chemistry","score_opus":0.029315675225262697,"score_gpt":0.36081865934085644,"score_spread":0.33150298411559376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411271765","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.101384595,0.00373993,0.85133576,0.0027371529,0.0006609901,0.0009336811,0.0013724996,0.030741509,0.0070938445],"genre_scores_gemma":[0.55953336,0.0011674115,0.41816518,0.002089105,0.0003246108,0.0006761387,0.0055107,0.0011802904,0.011353188],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9912634,0.003267895,0.0006107691,0.002527557,0.0018194005,0.00051093614],"domain_scores_gemma":[0.9646307,0.018803164,0.0031153017,0.0055098813,0.006991697,0.0009492384],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007190101,0.0022570493,0.0015356672,0.0028636164,0.001176998,0.0021913897,0.003078289,0.0024951547,0.0040965215],"category_scores_gemma":[0.03449277,0.0006342445,0.0011333838,0.0019785922,0.001631938,0.005477851,0.005699986,0.0045355977,0.0031181655],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005725286,0.0005212832,0.003342048,0.00074621284,0.00011936838,0.00017490567,0.00058057,0.026171446,0.029028881,0.0028366758,0.017961688,0.9179444],"study_design_scores_gemma":[0.00022332593,0.0006948025,0.0038585495,0.00023119427,0.00019026843,0.0003810822,0.00030297815,0.8805862,0.06914338,0.023899687,0.020322949,0.00016562195],"about_ca_topic_score_codex":0.00734252,"about_ca_topic_score_gemma":0.013693387,"teacher_disagreement_score":0.00734252,"about_ca_system_score_codex":0.0014700363,"about_ca_system_score_gemma":0.004604523,"threshold_uncertainty_score":0.03802532},"labels":[],"label_agreement":null},{"id":"W4411272126","doi":"10.1109/icse-companion66252.2025.00031","title":"Studying and Improving Code Understandability Through Atoms of Confusion","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Confusion; Computer science; Programming language; Code (set theory); Psychology","score_opus":0.031415400926340216,"score_gpt":0.30096577322633733,"score_spread":0.2695503722999971,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411272126","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9407057,0.00071257656,0.052062318,0.0012605853,0.000023220231,0.00012704436,0.00006391525,0.00039132533,0.004653294],"genre_scores_gemma":[0.97404003,0.00024940437,0.024663877,0.00013398872,0.000011179116,0.000074720985,0.000100369296,0.00008774777,0.0006387818],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.98698956,0.006060137,0.000690564,0.0014898313,0.004160478,0.00060945115],"domain_scores_gemma":[0.86037797,0.09746818,0.023345979,0.0056041037,0.010951822,0.0022519853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0115823345,0.0009675514,0.0005502882,0.0036770098,0.0009864081,0.0046490314,0.0010890394,0.0011543927,0.0011654372],"category_scores_gemma":[0.108239174,0.000497314,0.0005763741,0.0014104103,0.0026205212,0.008434713,0.0038370774,0.0021038635,0.00022467627],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004690818,0.00086458307,0.40183413,0.0011413977,0.00020045482,0.0005090257,0.1714805,0.0046213893,0.0277172,0.012805762,0.0018165187,0.37653998],"study_design_scores_gemma":[0.0001295708,0.002521244,0.6636956,0.0017106975,0.000771419,0.0016232108,0.096904896,0.08057818,0.04360966,0.07295076,0.034851626,0.00065319863],"about_ca_topic_score_codex":0.003731045,"about_ca_topic_score_gemma":0.0051902053,"teacher_disagreement_score":0.0115823345,"about_ca_system_score_codex":0.0028235444,"about_ca_system_score_gemma":0.0030418683,"threshold_uncertainty_score":0.061253965},"labels":[],"label_agreement":null},{"id":"W4411272370","doi":"10.1109/msr66628.2025.00067","title":"Understanding the Popularity of Packages in Maven Ecosystem","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Popularity; Ecosystem; Computer science; Ecology; Psychology","score_opus":0.0670270302834067,"score_gpt":0.29280575299561895,"score_spread":0.22577872271221225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411272370","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98627347,0.0009383711,0.00522275,0.00017927068,0.000020064454,0.000023810204,0.0032162883,0.0008654157,0.0032605084],"genre_scores_gemma":[0.97389776,0.00058012275,0.008277849,0.000075744865,0.000023780462,0.00007722635,0.014733575,0.00072995346,0.0016040503],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99837637,0.0004072656,0.00015428271,0.0004349041,0.0004502453,0.00017699518],"domain_scores_gemma":[0.9889558,0.0047230064,0.0024382635,0.0013148873,0.0020488647,0.00051918306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023499394,0.0005300425,0.00049081974,0.007458956,0.00048079513,0.0025979597,0.0007323899,0.0005100617,0.00091184926],"category_scores_gemma":[0.018929483,0.0003613766,0.0006527295,0.0057038893,0.00049760466,0.005250178,0.0016392134,0.0005074878,0.0007031306],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031223628,0.00010203007,0.9025049,0.00040384612,0.00020583259,0.0006362829,0.0022395845,0.009218572,0.0067110397,0.0030617958,0.008616735,0.065987095],"study_design_scores_gemma":[0.00002752306,0.00016893767,0.8573945,0.0001349204,0.00016269679,0.0018287635,0.0029684808,0.08641314,0.0040201214,0.005864924,0.040900584,0.000115409464],"about_ca_topic_score_codex":0.0045226724,"about_ca_topic_score_gemma":0.007436133,"teacher_disagreement_score":0.007458956,"about_ca_system_score_codex":0.0006452845,"about_ca_system_score_gemma":0.00046816742,"threshold_uncertainty_score":0.012427807},"labels":[],"label_agreement":null},{"id":"W4411272387","doi":"10.1109/msr66628.2025.00117","title":"Smells-sus: Sustainability Smells in IaC","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Sustainability; Computer science; Code smell; Ecology; Programming language; Software; Software quality; Biology; Software development","score_opus":0.006284715078031602,"score_gpt":0.2841022955910687,"score_spread":0.2778175805130371,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411272387","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9758302,0.00059739675,0.013182025,0.0011833865,0.000038080172,0.00016351088,0.0021804923,0.0015107446,0.0053142207],"genre_scores_gemma":[0.98194516,0.00026391828,0.012645718,0.00034614847,0.000017912567,0.00014309413,0.0029132634,0.00049167906,0.0012331705],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99085325,0.0033120965,0.0010547279,0.0010761127,0.0031233723,0.00058044377],"domain_scores_gemma":[0.9025258,0.055436287,0.021924865,0.008779001,0.009487174,0.0018468703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00840656,0.00060115114,0.00034197723,0.0041908356,0.0012722579,0.0025529577,0.00079529214,0.0008837012,0.0013489068],"category_scores_gemma":[0.07215974,0.0004117476,0.0004237302,0.0060681594,0.00237878,0.0048973076,0.003074418,0.0010900575,0.000381744],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029712362,0.00020004032,0.73843825,0.0015793337,0.00008392678,0.0025527794,0.081394374,0.0035989804,0.0104090255,0.009010816,0.015693538,0.13674177],"study_design_scores_gemma":[0.000039712362,0.0004811407,0.6694459,0.002266537,0.00011139891,0.0068721464,0.11128436,0.03093158,0.015408167,0.016101426,0.14667916,0.00037848778],"about_ca_topic_score_codex":0.0070462087,"about_ca_topic_score_gemma":0.012323075,"teacher_disagreement_score":0.00840656,"about_ca_system_score_codex":0.0023715138,"about_ca_system_score_gemma":0.002148245,"threshold_uncertainty_score":0.044458687},"labels":[],"label_agreement":null},{"id":"W4411337744","doi":"10.1109/designing66910.2025.00007","title":"Assessing LLMs for Front-end Software Architecture Knowledge","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Front and back ends; Computer science; Architecture; Software; Software engineering; Computer architecture; Operating system; History","score_opus":0.02297468073865381,"score_gpt":0.32742106027170753,"score_spread":0.3044463795330537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411337744","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8924518,0.0004364335,0.09079755,0.0005168132,0.00002848144,0.00059139763,0.0006894796,0.0023654203,0.012122512],"genre_scores_gemma":[0.92827404,0.00008816978,0.06989015,0.00005728646,0.000005278339,0.00027817595,0.00071636384,0.000052175634,0.00063839287],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9895628,0.0049161306,0.001154146,0.0011919978,0.0028036484,0.0003712136],"domain_scores_gemma":[0.8795103,0.09787039,0.0065164496,0.00763732,0.0063972184,0.0020684265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015428919,0.0009383426,0.00046170555,0.0042480286,0.00055510015,0.003636536,0.0012143516,0.0016744783,0.0026151799],"category_scores_gemma":[0.11577532,0.00041204956,0.00066350034,0.0014324636,0.0010743625,0.007458733,0.003978769,0.0012337944,0.0007926464],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017427718,0.0016979999,0.22394416,0.0015276135,0.00035439813,0.00023204967,0.00804021,0.052212477,0.019336322,0.012636619,0.0032094151,0.67506593],"study_design_scores_gemma":[0.00016519497,0.0038370122,0.15328524,0.0006287426,0.0002947546,0.00052393565,0.007746281,0.7459889,0.046543855,0.029950963,0.010790949,0.00024406878],"about_ca_topic_score_codex":0.0037031407,"about_ca_topic_score_gemma":0.0038507604,"teacher_disagreement_score":0.015428919,"about_ca_system_score_codex":0.0018567356,"about_ca_system_score_gemma":0.0018089681,"threshold_uncertainty_score":0.08159685},"labels":[],"label_agreement":null},{"id":"W4411379396","doi":"10.1007/978-981-96-8186-0_21","title":"Analogy-Augmented Generation with Procedural Memory for Procedure Generation","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Interface Biologics (Canada); Thales (Canada); Université de Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Analogy; Programming language; Theoretical computer science; Parallel computing","score_opus":0.021241273289699405,"score_gpt":0.2612090169738959,"score_spread":0.23996774368419652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411379396","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005518022,0.00013238736,0.97065026,0.000070065355,0.000091405105,0.000059665683,0.00006779383,0.004704656,0.018705746],"genre_scores_gemma":[0.21770926,0.00013213656,0.76155794,0.00012297058,0.00005314766,0.00014979618,0.0003297322,0.0013017516,0.018643254],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993625,0.00019653133,0.000036627476,0.00012504843,0.00021678628,0.00006248258],"domain_scores_gemma":[0.9990175,0.0004416044,0.00003094764,0.0003969523,0.000092038696,0.000020906937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00053212326,0.0006151027,0.00061570125,0.00052022986,0.0006029946,0.0011271831,0.0022233613,0.0011288121,0.022038471],"category_scores_gemma":[0.002645329,0.00053152384,0.0009819382,0.00066851656,0.00095430383,0.0020542764,0.0018213167,0.0018282144,0.0042127566],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028817946,0.00014207834,0.00025901114,0.00026582793,0.000033737102,0.00033788738,0.00042394517,0.04228471,0.024265727,0.4391528,0.011562143,0.480984],"study_design_scores_gemma":[0.000095817006,0.00013785546,0.00026365826,0.00006054241,0.0000777546,0.00060279673,0.00007193533,0.5212413,0.041429646,0.37604633,0.059905067,0.00006729248],"about_ca_topic_score_codex":0.00090583606,"about_ca_topic_score_gemma":0.0012530402,"teacher_disagreement_score":0.022038471,"about_ca_system_score_codex":0.00045113193,"about_ca_system_score_gemma":0.00061491807,"threshold_uncertainty_score":0.073726},"labels":[],"label_agreement":null},{"id":"W4411449687","doi":"10.1145/3715730","title":"Towards Diverse Program Transformations for Program Simplification","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Concordia University","funders":"Engineering and Physical Sciences Research Council; Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Program comprehension; Computer science; Source lines of code; Program transformation; Maintainability; Program slicing; Programming language; Heuristics; Software engineering; Set (abstract data type); Code (set theory); Static program analysis; Program analysis; Software; Source code; Software maintenance; Software quality; Software system; Software development; Operating system","score_opus":0.01895505048035933,"score_gpt":0.29468973969045453,"score_spread":0.2757346892100952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411449687","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07891822,0.0007444045,0.8396897,0.00067746977,0.000076830656,0.0005430279,0.0019205789,0.07463872,0.002791066],"genre_scores_gemma":[0.16293325,0.0005369023,0.8156297,0.00043749553,0.00004907389,0.000443658,0.011663845,0.0066437284,0.0016623401],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9915176,0.0024091462,0.0007811349,0.0020944222,0.0028097387,0.00038799507],"domain_scores_gemma":[0.9787411,0.009060334,0.0018319228,0.006939915,0.0031377694,0.00028895529],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004671357,0.0020068905,0.0011746656,0.0049992283,0.0009833543,0.0019444107,0.0025175272,0.001193202,0.0018254946],"category_scores_gemma":[0.031859547,0.0012907973,0.0029241976,0.003388688,0.0015068981,0.0036079888,0.0038227183,0.0030333325,0.0020383922],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059569813,0.00079588726,0.033832572,0.0019375829,0.0003955423,0.00092092477,0.002834933,0.04990559,0.08773897,0.015261948,0.029888513,0.7758919],"study_design_scores_gemma":[0.00032165158,0.0007008758,0.014466686,0.00050207373,0.00046942977,0.0019399015,0.0010244648,0.6992407,0.13642214,0.04941817,0.095317505,0.00017639733],"about_ca_topic_score_codex":0.0024790931,"about_ca_topic_score_gemma":0.0061790803,"teacher_disagreement_score":0.0049992283,"about_ca_system_score_codex":0.0010069477,"about_ca_system_score_gemma":0.0026858693,"threshold_uncertainty_score":0.024704754},"labels":[],"label_agreement":null},{"id":"W4411449706","doi":"10.1145/3715738","title":"Code Change Intention, Development Artifact, and History Vulnerability: Putting Them Together for Vulnerability Fix Detection by LLM","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); University of Manitoba","funders":"","keywords":"Computer science; Vulnerability (computing); Artifact (error); Commit; Context (archaeology); Leverage (statistics); Computer security; Vulnerability assessment; Data science; Artificial intelligence; Database; Psychology","score_opus":0.038219735064336666,"score_gpt":0.2594014123250999,"score_spread":0.2211816772607632,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411449706","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1830833,0.0029068182,0.70104325,0.0018807795,0.00026293017,0.00047720643,0.005141542,0.101660326,0.003543777],"genre_scores_gemma":[0.5136851,0.00039597557,0.4736831,0.00052343484,0.00005220999,0.00025833905,0.008393742,0.0009807955,0.002027283],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99719703,0.00085618434,0.00026203576,0.00087985065,0.0006101724,0.00019469908],"domain_scores_gemma":[0.9921863,0.0051772282,0.000526986,0.0010179378,0.000851364,0.00024019321],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026221853,0.0020513244,0.00065315387,0.0030863832,0.0005843059,0.0015358954,0.0020473455,0.0016244845,0.002722497],"category_scores_gemma":[0.012460486,0.0007147497,0.0017654208,0.0010426735,0.000717176,0.0035313678,0.0024375673,0.003296476,0.001783106],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006594124,0.0007115697,0.036223132,0.001098205,0.00026742942,0.0007885565,0.0013485575,0.078948855,0.032424204,0.0037419589,0.016732262,0.8270559],"study_design_scores_gemma":[0.000053135318,0.00015599193,0.0041927085,0.00007858465,0.00009801596,0.00025201682,0.00032006032,0.9633134,0.017162412,0.007351804,0.006948345,0.00007357218],"about_ca_topic_score_codex":0.013034832,"about_ca_topic_score_gemma":0.027008727,"teacher_disagreement_score":0.013034832,"about_ca_system_score_codex":0.0015424576,"about_ca_system_score_gemma":0.0025147721,"threshold_uncertainty_score":0.025917888},"labels":[],"label_agreement":null},{"id":"W4411449758","doi":"10.1145/3715734","title":"An Empirical Study on Release-Wise Refactoring Patterns","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Code refactoring; Computer science; Cohesion (chemistry); Quality (philosophy); Software deployment; Java; Code (set theory); Software evolution; Software engineering; Software; Programming language; Software system","score_opus":0.023231487738195986,"score_gpt":0.30415486897821586,"score_spread":0.28092338124001986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411449758","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99660265,0.000233511,0.0016347193,0.00013753073,0.0000067499404,0.0000699609,0.00025186728,0.0000383194,0.0010245097],"genre_scores_gemma":[0.9970715,0.0001469436,0.0017875796,0.000043634656,0.000007413488,0.00009229391,0.00038429344,0.000030411971,0.00043590713],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9879704,0.003681946,0.0016893588,0.0019486428,0.004063808,0.0006457414],"domain_scores_gemma":[0.6550581,0.20110068,0.09564765,0.011923021,0.03135919,0.004911321],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013191582,0.000385818,0.00032427444,0.0027133373,0.00071452634,0.0019519653,0.0011248356,0.0007648472,0.0013374736],"category_scores_gemma":[0.11947806,0.0005055139,0.00044012794,0.0034217234,0.0012370981,0.002800369,0.0012717552,0.0015443306,0.00041822914],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002143703,0.00024987536,0.9615933,0.00021691235,0.00007980589,0.00025676252,0.008212707,0.0005025575,0.0014705703,0.00035583443,0.0005192449,0.026328055],"study_design_scores_gemma":[0.0000158905,0.00033806317,0.9877395,0.000074194104,0.000030803985,0.00031340602,0.0065251254,0.0021522033,0.0006351253,0.00026513005,0.0018806283,0.000029827272],"about_ca_topic_score_codex":0.0024933196,"about_ca_topic_score_gemma":0.0036969848,"teacher_disagreement_score":0.013191582,"about_ca_system_score_codex":0.0011455609,"about_ca_system_score_gemma":0.0011218004,"threshold_uncertainty_score":0.069764555},"labels":[],"label_agreement":null},{"id":"W4411449952","doi":"10.1145/3715729","title":"An Empirical Study of Suppressed Static Analysis Warnings","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Spectrum analyzer; False positive paradox; Python (programming language); Computer science; Software; Static analysis; Scalability; Empirical research; False positives and false negatives; Code (set theory); Artificial intelligence; Programming language; Statistics; Telecommunications; Operating system; Mathematics","score_opus":0.0138911653291388,"score_gpt":0.29776158001213404,"score_spread":0.28387041468299523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411449952","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9968978,0.00018848245,0.0012592062,0.0001713766,0.000009463416,0.000046594963,0.0001485798,0.000074419135,0.0012041285],"genre_scores_gemma":[0.9976841,0.00012762762,0.0012596372,0.00008177973,0.000012436458,0.000077315184,0.0002744886,0.00004030961,0.00044215942],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9807842,0.008362874,0.002045238,0.0023499785,0.0056635486,0.00079416466],"domain_scores_gemma":[0.6027452,0.23665886,0.094995245,0.018203497,0.041905228,0.0054919543],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019658288,0.00053117843,0.00039125548,0.0034282645,0.0010288821,0.0016436818,0.0012885832,0.00088185864,0.0013904454],"category_scores_gemma":[0.18644525,0.00058365555,0.00031004156,0.0025687465,0.0019270842,0.003912123,0.0020740998,0.0021394193,0.00042015786],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036868837,0.00045437046,0.90534794,0.00045970382,0.00010193359,0.0006015427,0.038610827,0.000521938,0.0026438616,0.0006614327,0.0017128662,0.04851492],"study_design_scores_gemma":[0.00004510229,0.0011233039,0.9454249,0.00035198082,0.00007777489,0.0012527549,0.030822007,0.00671124,0.0025591368,0.00095570274,0.010582405,0.00009365995],"about_ca_topic_score_codex":0.0019505593,"about_ca_topic_score_gemma":0.0023237583,"teacher_disagreement_score":0.019658288,"about_ca_system_score_codex":0.0006959908,"about_ca_system_score_gemma":0.0012327883,"threshold_uncertainty_score":0.10396415},"labels":[],"label_agreement":null},{"id":"W4411450091","doi":"10.1145/3715724","title":"CKTyper: Enhancing Type Inference for Java Code Snippets by Leveraging Crowdsourcing Knowledge in Stack Overflow","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Basic and Applied Basic Research Foundation of Guangdong Province","keywords":"Snippet; Computer science; Crowdsourcing; Context (archaeology); Code (set theory); Inference; Set (abstract data type); Information retrieval; Type inference; Java; Source code; World Wide Web; Artificial intelligence; Programming language","score_opus":0.017235033584320957,"score_gpt":0.2818270580510308,"score_spread":0.2645920244667099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411450091","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08997626,0.00166789,0.80845,0.002548062,0.0008141303,0.0016208582,0.020026205,0.06303087,0.011865732],"genre_scores_gemma":[0.33157668,0.00072578294,0.61446005,0.001685774,0.00041396834,0.0015193606,0.03289322,0.0043893754,0.012335844],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99311054,0.0014570398,0.00043497572,0.002185591,0.0024254338,0.0003864954],"domain_scores_gemma":[0.9862316,0.007161349,0.0011921304,0.0028031862,0.0021257738,0.00048598484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004473238,0.0027966776,0.0012958692,0.0086967535,0.0021100813,0.0021237785,0.0029533668,0.0026795364,0.0041013528],"category_scores_gemma":[0.03155835,0.0008197028,0.002352061,0.0034475978,0.0016158328,0.0055519072,0.0049837944,0.0028493935,0.0029160038],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015639989,0.0007611581,0.027923858,0.0028117786,0.00048161994,0.0030401805,0.005110279,0.05347438,0.04961157,0.014269745,0.08702515,0.75392634],"study_design_scores_gemma":[0.0002922559,0.0002604919,0.015737887,0.00044509314,0.00029029336,0.0010085675,0.0019749457,0.76394,0.05160142,0.058879565,0.105106294,0.00046321267],"about_ca_topic_score_codex":0.026329812,"about_ca_topic_score_gemma":0.05241246,"teacher_disagreement_score":0.026329812,"about_ca_system_score_codex":0.0020531046,"about_ca_system_score_gemma":0.0037701996,"threshold_uncertainty_score":0.052353144},"labels":[],"label_agreement":null},{"id":"W4411450398","doi":"10.1145/3729383","title":"Automated Extraction and Analysis of Developer's Rationale in Open Source Software","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal","funders":"","keywords":"Computer science; Software engineering; Software; Open source; Generalization; Open source software; Data science; Artificial intelligence; Programming language","score_opus":0.015592076908858506,"score_gpt":0.27953756571677185,"score_spread":0.26394548880791335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411450398","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12657768,0.00067045074,0.8495694,0.0018354555,0.00012375254,0.0009961085,0.002986051,0.013092906,0.0041483273],"genre_scores_gemma":[0.2689251,0.00032568575,0.72207373,0.00018695078,0.00005216657,0.0003411942,0.005399452,0.00085066666,0.0018449313],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98792946,0.0049188524,0.0010640712,0.0010860099,0.004616786,0.0003847305],"domain_scores_gemma":[0.94104934,0.038157392,0.006433772,0.0052429046,0.00860747,0.00050914765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012047728,0.0012338902,0.0006590258,0.010568864,0.0016689316,0.0037571178,0.0015421441,0.0016270586,0.0014001096],"category_scores_gemma":[0.052872553,0.0011308952,0.0014092717,0.0033692087,0.0011086505,0.005441541,0.0030662783,0.0024517085,0.0008733144],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027110858,0.0005134546,0.042458232,0.0015187131,0.00019671253,0.002444146,0.010098983,0.020425035,0.046405468,0.035266623,0.017582603,0.82281893],"study_design_scores_gemma":[0.0001687259,0.0003478054,0.03390755,0.0011938318,0.0003219292,0.0017818331,0.006127339,0.6698895,0.09076859,0.10928236,0.08572871,0.00048184037],"about_ca_topic_score_codex":0.0043702824,"about_ca_topic_score_gemma":0.013126711,"teacher_disagreement_score":0.012047728,"about_ca_system_score_codex":0.0015387858,"about_ca_system_score_gemma":0.0046923556,"threshold_uncertainty_score":0.06371522},"labels":[],"label_agreement":null},{"id":"W4411488587","doi":"10.1145/3744920","title":"On the Utility of Domain Modeling Assistance with Large Language Models","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Université de Montréal","funders":"","keywords":"Computer science; Modeling language; Domain (mathematical analysis); Usability; Software engineering; Domain analysis; Abstraction; Domain-specific language; Process (computing); Model-driven architecture; Context (archaeology); Subject-matter expert; Domain model; Software; Domain engineering; Human–computer interaction; Software development; Data science; Artificial intelligence; Domain knowledge; Programming language; Component-based software engineering; Software construction","score_opus":0.06253040734009754,"score_gpt":0.3178245270606738,"score_spread":0.25529411972057625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411488587","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45064798,0.0013483646,0.5214788,0.0023520573,0.00010345089,0.000563958,0.0005417783,0.011250888,0.011712744],"genre_scores_gemma":[0.7524036,0.00040475538,0.24437149,0.00026597327,0.000025179295,0.00013744067,0.00044805047,0.0004936112,0.0014497867],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9906313,0.006776527,0.00031402457,0.0009797461,0.0011348985,0.0001635427],"domain_scores_gemma":[0.8425185,0.1450806,0.0014694278,0.0066917525,0.0035120754,0.0007277401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011555801,0.0013556592,0.0005168503,0.0013824096,0.00085046847,0.0025854092,0.0018226053,0.001993002,0.0031510289],"category_scores_gemma":[0.099085785,0.00064783957,0.0005574376,0.00088639295,0.0011099064,0.007871268,0.0028971694,0.0020349913,0.0009415102],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028924793,0.002324549,0.015331045,0.0013024153,0.00027184005,0.0006448142,0.005812002,0.16960584,0.033127345,0.014728813,0.0071411985,0.74681777],"study_design_scores_gemma":[0.00015648104,0.0010008528,0.0032573417,0.00017650609,0.00010842248,0.00031731068,0.001188253,0.9523172,0.020741228,0.012933552,0.007708208,0.000094733674],"about_ca_topic_score_codex":0.0069232886,"about_ca_topic_score_gemma":0.008747774,"teacher_disagreement_score":0.011555801,"about_ca_system_score_codex":0.0009474091,"about_ca_system_score_gemma":0.0014880926,"threshold_uncertainty_score":0.061113656},"labels":[],"label_agreement":null},{"id":"W4411489800","doi":"10.5753/cibse.2025.35294","title":"Generación de Pruebas Unitarias con LLMs en Entornos Industriales: Desafíos, Evolución y Lecciones Prácticas","year":2025,"lang":"es","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nexen (Canada)","funders":"","keywords":"Humanities; Political science; Philosophy","score_opus":0.023077582971818097,"score_gpt":0.3146113494746925,"score_spread":0.2915337665028744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411489800","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5316622,0.0018009901,0.43219927,0.00072728493,0.00013483346,0.0005661606,0.00061302615,0.013890838,0.018405417],"genre_scores_gemma":[0.7638338,0.0008889988,0.22495593,0.0001130324,0.000025607413,0.00038622404,0.0006278156,0.0011836418,0.007984967],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968629,0.0007961421,0.00022347721,0.0005807882,0.0012510741,0.00028553637],"domain_scores_gemma":[0.98500293,0.007286595,0.0011624685,0.0032087113,0.0028018442,0.0005374874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047471793,0.0012920592,0.0007695547,0.002077478,0.00074612425,0.0039971066,0.0017890809,0.0010983538,0.0058501386],"category_scores_gemma":[0.019423934,0.00088142225,0.0012420026,0.0014906235,0.0009802978,0.0043853167,0.0020932017,0.0015872462,0.0015241462],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015325804,0.0010325136,0.052304987,0.0015604561,0.0002248507,0.0005122092,0.0058608535,0.19984362,0.06698565,0.017959695,0.0033780735,0.64880455],"study_design_scores_gemma":[0.00024310606,0.0023948636,0.03610318,0.0007792329,0.0006422755,0.000607189,0.004293486,0.7506496,0.12740639,0.021097194,0.055530906,0.00025262017],"about_ca_topic_score_codex":0.009012172,"about_ca_topic_score_gemma":0.008915773,"teacher_disagreement_score":0.009012172,"about_ca_system_score_codex":0.0018431935,"about_ca_system_score_gemma":0.0027869656,"threshold_uncertainty_score":0.025105774},"labels":[],"label_agreement":null},{"id":"W4411523084","doi":"10.1145/3728947","title":"The First Prompt Counts the Most! An Evaluation of Large Language Models on Iterative Example-Based Code Generation","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Code (set theory); Computer science; Benchmark (surveying); Iterative and incremental development; Code generation; Process (computing); Natural language generation; Natural language; Programming language; Software engineering; Artificial intelligence; Computer security; Geography","score_opus":0.03728653452721785,"score_gpt":0.29708271316430906,"score_spread":0.2597961786370912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411523084","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6797331,0.0032200613,0.26637298,0.0026740246,0.00041219144,0.0008985984,0.0021349262,0.024596568,0.019957498],"genre_scores_gemma":[0.6968561,0.00065398787,0.2905954,0.00051459094,0.00004288229,0.0004022537,0.0056602657,0.002079155,0.0031953412],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9839138,0.008508175,0.0007676229,0.0014534666,0.004884723,0.0004721425],"domain_scores_gemma":[0.9319447,0.045480896,0.0021422768,0.010447636,0.008788329,0.0011961933],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010639799,0.0013957868,0.00089741754,0.0013409805,0.0005595933,0.002291641,0.0024234543,0.0015581814,0.0037069446],"category_scores_gemma":[0.074627034,0.0004931991,0.0010555172,0.0010932882,0.0012826596,0.0041286475,0.002586553,0.0021700477,0.0014507172],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032868145,0.0022746332,0.024948789,0.0030347393,0.00048793328,0.00059938116,0.0027288178,0.20937644,0.03705904,0.016428515,0.031635907,0.6681389],"study_design_scores_gemma":[0.00048029586,0.0025031138,0.008158939,0.00045163042,0.0002104257,0.00043140783,0.001215738,0.8979132,0.046336547,0.010256997,0.031866904,0.00017473844],"about_ca_topic_score_codex":0.0044884337,"about_ca_topic_score_gemma":0.0053305365,"teacher_disagreement_score":0.010639799,"about_ca_system_score_codex":0.0014666169,"about_ca_system_score_gemma":0.00235127,"threshold_uncertainty_score":0.056269348},"labels":[],"label_agreement":null},{"id":"W4411523132","doi":"10.1145/3728931","title":"Understanding Practitioners’ Expectations on Clear Code Review Comments","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"CLARITY; Computer science; Constructive; Relevance (law); Set (abstract data type); Code (set theory); Process (computing); Data science; Code review; Programming language; Software; Software quality; Political science; Software development","score_opus":0.058257941025824454,"score_gpt":0.2989452171374843,"score_spread":0.24068727611165983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411523132","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7414639,0.0048350366,0.1850963,0.028308269,0.0008119354,0.0027745187,0.0010313868,0.004377065,0.03130161],"genre_scores_gemma":[0.9110454,0.0013644992,0.07729882,0.0037580007,0.00025619654,0.0013740979,0.0010332373,0.0007037598,0.0031659126],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.7449617,0.14688882,0.021116758,0.011231133,0.07154385,0.004257803],"domain_scores_gemma":[0.15616548,0.5341199,0.06310776,0.020182382,0.21961108,0.0068133534],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.19054729,0.0008006697,0.0008658379,0.008477343,0.0024751602,0.008369639,0.0021982556,0.0032094927,0.0025632652],"category_scores_gemma":[0.66641456,0.00095256424,0.0008508506,0.0034167124,0.0030519264,0.010792756,0.0054544252,0.0033987323,0.0022070475],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014045611,0.00045554998,0.1496185,0.007530111,0.00023845574,0.0011584088,0.17607987,0.004122589,0.0352952,0.01355615,0.043158352,0.56738234],"study_design_scores_gemma":[0.0007760562,0.0033703805,0.27468917,0.016904246,0.0007961949,0.0042357896,0.18924725,0.07326375,0.05308099,0.0470354,0.33487543,0.0017253795],"about_ca_topic_score_codex":0.0038341694,"about_ca_topic_score_gemma":0.004547938,"teacher_disagreement_score":0.19054729,"about_ca_system_score_codex":0.007260213,"about_ca_system_score_gemma":0.010797898,"threshold_uncertainty_score":0.9981993},"labels":[],"label_agreement":null},{"id":"W4411534996","doi":"10.1007/978-3-031-96590-6_4","title":"Operating Under Constraints: Identifying Requirements for Enhanced Cyber Resilience Management","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Resilience (materials science); Computer science; Computer security","score_opus":0.03145355182758525,"score_gpt":0.31399095957602985,"score_spread":0.2825374077484446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411534996","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18865964,0.0002911009,0.7626227,0.001288575,0.0001388654,0.0005172317,0.0007412193,0.0011705942,0.04457013],"genre_scores_gemma":[0.92044204,0.0001487401,0.07571772,0.000086718064,0.00003945688,0.00017776956,0.00035994197,0.00019620013,0.0028313238],"study_design_codex":"simulation_or_modeling","study_design_gemma":"qualitative","domain_scores_codex":[0.9981857,0.00030461364,0.00010394731,0.00028262092,0.00066870183,0.00045427485],"domain_scores_gemma":[0.99087065,0.0046611372,0.0016068337,0.0011156087,0.0012434619,0.0005022895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016838097,0.0009670463,0.00075113203,0.0008843561,0.0006170307,0.0027919288,0.0018765265,0.0010868991,0.0069065983],"category_scores_gemma":[0.013491394,0.0004452492,0.0005727406,0.00085325015,0.0014585659,0.0054411325,0.002125849,0.0014944394,0.0006754781],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050307607,0.00026289898,0.0065890057,0.00063756737,0.00007206929,0.0010018855,0.000574512,0.64245,0.02517624,0.21840152,0.0059377756,0.09839352],"study_design_scores_gemma":[0.000017775586,0.00013456284,0.002379837,0.000078067125,0.000030084038,0.0002765063,0.0005414766,0.89440334,0.008129379,0.09002529,0.0039451234,0.000038432965],"about_ca_topic_score_codex":0.002591561,"about_ca_topic_score_gemma":0.002899319,"teacher_disagreement_score":0.0069065983,"about_ca_system_score_codex":0.0009145409,"about_ca_system_score_gemma":0.0018059983,"threshold_uncertainty_score":0.023104846},"labels":[],"label_agreement":null},{"id":"W4411552659","doi":"10.1109/icse55347.2025.00012","title":"Reasoning Runtime Behavior of a Program with LLM: How Far are We?","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Science Foundation of Ningbo","keywords":"Computer science; Programming language","score_opus":0.016730493101346253,"score_gpt":0.28722802632078254,"score_spread":0.27049753321943626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411552659","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40976247,0.0057303477,0.5240401,0.009652192,0.00014649137,0.000267992,0.0019366086,0.042005934,0.006457874],"genre_scores_gemma":[0.675873,0.0011167201,0.31641257,0.0009813998,0.00006097524,0.00010090452,0.0027401152,0.0016609086,0.0010533382],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99073255,0.0035113662,0.00066962186,0.002157777,0.00237191,0.00055670086],"domain_scores_gemma":[0.96647036,0.016492326,0.0038123268,0.009044179,0.0032752624,0.0009055],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00960756,0.0014034603,0.0013123457,0.0020110675,0.0009542106,0.004675909,0.0035317773,0.0021117658,0.0015234591],"category_scores_gemma":[0.051792473,0.00067800726,0.0017235054,0.001610434,0.0022248926,0.011327566,0.003253419,0.002864101,0.0010275944],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011734299,0.00075644715,0.08402202,0.0014774818,0.0004276719,0.00027810095,0.0022202441,0.08329098,0.022646248,0.015239308,0.009082207,0.77938586],"study_design_scores_gemma":[0.000108816734,0.0006907748,0.022519806,0.0005514219,0.0004297949,0.00044293294,0.0024711601,0.86834586,0.045636512,0.043926552,0.014637951,0.00023846544],"about_ca_topic_score_codex":0.019390525,"about_ca_topic_score_gemma":0.020750107,"teacher_disagreement_score":0.019390525,"about_ca_system_score_codex":0.0020762624,"about_ca_system_score_gemma":0.00451365,"threshold_uncertainty_score":0.050810218},"labels":[],"label_agreement":null},{"id":"W4411617963","doi":"10.51847/2lybtjlmml","title":"10.51847/2LyBtjLMml","year":2000,"lang":"en","type":"article","venue":"Time to knit","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Estimation; Software; Reliability engineering; Engineering; Systems engineering; Operating system","score_opus":0.00814562802463771,"score_gpt":0.20305294405862542,"score_spread":0.1949073160339877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411617963","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0072498024,0.0011091854,0.016180532,0.0019984625,0.0011829747,0.00020030246,0.0022548311,0.0054291133,0.9643947],"genre_scores_gemma":[0.012477208,0.0003977,0.005055784,0.00044107318,0.000111236724,0.00009747899,0.0016001936,0.00059136597,0.97922796],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99944645,0.00007927971,0.000042327774,0.0001729809,0.00016856803,0.000090370406],"domain_scores_gemma":[0.9990238,0.00026552868,0.000057052974,0.00017173078,0.00032490958,0.00015703871],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0007077621,0.0013627029,0.0006508621,0.0021193952,0.0009303822,0.0030662257,0.0010284078,0.0029706527,0.92600715],"category_scores_gemma":[0.0013335504,0.00040077494,0.00049423537,0.0020379897,0.0007701981,0.0018732919,0.0016913611,0.0012329826,0.9293292],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034425117,0.00028131183,0.0019184819,0.00043601298,0.000026313986,0.0004244461,0.0001229144,0.0010829653,0.009013134,0.0068693715,0.19892201,0.78055876],"study_design_scores_gemma":[0.00007884493,0.00011389392,0.0019183169,0.00025125398,0.000019034927,0.0005436104,0.00017682277,0.004098966,0.0029432445,0.0019071079,0.9879154,0.000033537002],"about_ca_topic_score_codex":0.0035284094,"about_ca_topic_score_gemma":0.0028868208,"teacher_disagreement_score":0.07399285,"about_ca_system_score_codex":0.0010479669,"about_ca_system_score_gemma":0.0005173711,"threshold_uncertainty_score":0.105541766},"labels":[],"label_agreement":null},{"id":"W4411617973","doi":"10.51847/gj5e8aptzj","title":"10.51847/gJ5E8aPTzj","year":2000,"lang":"en","type":"article","venue":"Time to knit","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Estimation; Computer science; Reliability engineering; Software; Statistics; Mathematics; Systems engineering; Engineering","score_opus":0.008316632172210351,"score_gpt":0.20291448068904194,"score_spread":0.1945978485168316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411617973","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077592917,0.00079226523,0.009411482,0.0020270026,0.0010242226,0.00022177021,0.0029141179,0.004668723,0.9711811],"genre_scores_gemma":[0.015561596,0.00043839612,0.0054901415,0.00052689743,0.00013170962,0.000100680816,0.0020761746,0.0007929294,0.97488153],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.999433,0.00007348336,0.000055932043,0.00014790815,0.00018655967,0.00010320121],"domain_scores_gemma":[0.9985153,0.00037043708,0.00010722233,0.00030776238,0.0004516649,0.00024763102],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0009474729,0.0010839811,0.0006085186,0.0026921812,0.0010122003,0.0032368721,0.0010477563,0.0032483449,0.9369919],"category_scores_gemma":[0.0020839255,0.00041147138,0.00048144226,0.002638997,0.00088026194,0.0022605506,0.001865861,0.000994829,0.9340549],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003412776,0.0002711869,0.003032844,0.000411961,0.000026192516,0.00045518676,0.0001714442,0.0007165595,0.0061598783,0.007172948,0.20536011,0.7758804],"study_design_scores_gemma":[0.00008919525,0.00011829621,0.0040739705,0.00030759937,0.000025947247,0.00067559787,0.00024608654,0.0018462889,0.002270499,0.0023581705,0.9879534,0.00003489883],"about_ca_topic_score_codex":0.0045597223,"about_ca_topic_score_gemma":0.0027650346,"teacher_disagreement_score":0.06300813,"about_ca_system_score_codex":0.00095777155,"about_ca_system_score_gemma":0.00057605223,"threshold_uncertainty_score":0.08987343},"labels":[],"label_agreement":null},{"id":"W4411950757","doi":"10.1109/forge66646.2025.00017","title":"Augmenting Large Language Models with Static Code Analysis for Automated Code Quality Improvements","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Programming language; Code (set theory); Code generation; Software quality; Static program analysis; KPI-driven code analysis; Software; Operating system; Software development","score_opus":0.026496059906458435,"score_gpt":0.3556332030345775,"score_spread":0.3291371431281191,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411950757","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.083903514,0.00040439118,0.880713,0.00082497566,0.0000719002,0.0004132982,0.0004522438,0.031082556,0.0021340444],"genre_scores_gemma":[0.33599886,0.00019049108,0.65963817,0.00022892612,0.000026188563,0.0002735227,0.0007627749,0.0019796635,0.00090146356],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9926778,0.0033638377,0.00045721995,0.0010746735,0.0022025334,0.00022397553],"domain_scores_gemma":[0.96301377,0.02101004,0.0034037163,0.007864109,0.004419868,0.00028838916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007984326,0.0015041314,0.00083167275,0.0023252754,0.00057431275,0.002950546,0.0021752343,0.0009735034,0.0017068572],"category_scores_gemma":[0.052005626,0.0011043288,0.0015652335,0.0012277026,0.0010120368,0.004765133,0.0031080819,0.0021425385,0.0007712793],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008459005,0.0009021932,0.024402635,0.0011580976,0.00029313395,0.00062689657,0.0045225746,0.24814089,0.047301725,0.015357987,0.007693419,0.64875454],"study_design_scores_gemma":[0.00008393222,0.00030699666,0.002282116,0.000115340685,0.00015875825,0.00016022404,0.00036764686,0.9498242,0.022741824,0.013563381,0.010298107,0.00009746973],"about_ca_topic_score_codex":0.0061371992,"about_ca_topic_score_gemma":0.011641176,"teacher_disagreement_score":0.007984326,"about_ca_system_score_codex":0.0014909932,"about_ca_system_score_gemma":0.0038487695,"threshold_uncertainty_score":0.04222566},"labels":[],"label_agreement":null},{"id":"W4412033360","doi":"10.1145/3747289","title":"Towards Understanding Refactoring Engine Bugs","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code refactoring; Computer science; Software engineering; Programming language; Software","score_opus":0.1393592182729809,"score_gpt":0.35114929550490953,"score_spread":0.21179007723192864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412033360","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6945721,0.010366413,0.28270212,0.0029608104,0.0001880183,0.0007966926,0.00179542,0.0024385096,0.004179936],"genre_scores_gemma":[0.6280869,0.0044283527,0.36142808,0.00097165385,0.00012677711,0.00036299968,0.0024827542,0.00066534115,0.0014470747],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98829275,0.002196435,0.0019588892,0.0023473313,0.004401916,0.000802729],"domain_scores_gemma":[0.91173893,0.042713024,0.020508911,0.0070188614,0.01715204,0.00086828636],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009957315,0.0020744433,0.0011150754,0.016714716,0.0012228668,0.0038483685,0.0026193922,0.0024647159,0.00097651547],"category_scores_gemma":[0.05783904,0.0011695626,0.0014947853,0.0053016106,0.0019954701,0.0108021675,0.0025798217,0.0026261278,0.00037527917],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025833596,0.0007576402,0.51327986,0.0036948689,0.00032875777,0.0035324234,0.024253055,0.0076661417,0.033870555,0.0134589,0.0032197512,0.39567968],"study_design_scores_gemma":[0.00021266192,0.0016520261,0.6461844,0.0052340925,0.0017689752,0.015066074,0.028654069,0.11480818,0.07111798,0.059866786,0.054769427,0.00066546747],"about_ca_topic_score_codex":0.0065541137,"about_ca_topic_score_gemma":0.0071457117,"teacher_disagreement_score":0.016714716,"about_ca_system_score_codex":0.0014112251,"about_ca_system_score_gemma":0.003249174,"threshold_uncertainty_score":0.05265993},"labels":[],"label_agreement":null},{"id":"W4412602677","doi":"10.1145/3749100","title":"<scp>AntiCopyPaster</scp> 3.0: Just-in-Time Clone Refactoring","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Computer science; clone (Java method); Programming language; Software engineering; Software; Biology","score_opus":0.06483042430187017,"score_gpt":0.328490238395962,"score_spread":0.2636598140940918,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412602677","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035675496,0.00023099309,0.15256935,0.0010198529,0.00033104362,0.0003785255,0.01191452,0.8219044,0.008083736],"genre_scores_gemma":[0.073038764,0.00066570676,0.45715266,0.0035291265,0.00046438316,0.001729306,0.08915731,0.3364901,0.037772577],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9947476,0.0007365485,0.00045567422,0.00095342187,0.0027846268,0.00032214943],"domain_scores_gemma":[0.9749341,0.008711547,0.0019534973,0.008686873,0.0047656456,0.00094835093],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004940955,0.0025532052,0.001048839,0.0030582498,0.0010569904,0.0027049824,0.0056706094,0.002675771,0.05724053],"category_scores_gemma":[0.029602423,0.002219395,0.0016375872,0.0027460952,0.0016084342,0.005712482,0.0047054575,0.0039203316,0.03906807],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008017121,0.00016940961,0.0030236016,0.0009711004,0.00012536449,0.0007683147,0.0006957206,0.0023084648,0.017160947,0.0041002952,0.80048484,0.16939011],"study_design_scores_gemma":[0.00073982985,0.0006290491,0.012410005,0.00051358726,0.00013154813,0.002433758,0.00017413174,0.093680285,0.08209113,0.017232865,0.78944814,0.0005157567],"about_ca_topic_score_codex":0.0085230125,"about_ca_topic_score_gemma":0.014602606,"teacher_disagreement_score":0.05724053,"about_ca_system_score_codex":0.0015698599,"about_ca_system_score_gemma":0.002866233,"threshold_uncertainty_score":0.19148862},"labels":[],"label_agreement":null},{"id":"W4412703947","doi":"10.1145/3696630.3731674","title":"An Empirical Study on the Impact of Gender Diversity on Code Quality in AI Systems","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Diversity (politics); Computer science; Code (set theory); Quality (philosophy); Gender diversity; Empirical research; Programming language; Mathematics; Sociology; Statistics; Business; Anthropology; Epistemology; Philosophy","score_opus":0.1336716054099922,"score_gpt":0.44618265928969353,"score_spread":0.31251105387970135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412703947","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99699855,0.00015598437,0.00063191145,0.00018969388,0.000006920005,0.00002513068,0.00009444235,0.000004373859,0.0018930912],"genre_scores_gemma":[0.99906355,0.000072613635,0.0003601371,0.000045919256,0.000008931625,0.000029988913,0.000054937613,0.0000060084076,0.0003579256],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9899429,0.00498781,0.0006269024,0.0010363085,0.0026696369,0.00073640456],"domain_scores_gemma":[0.7336745,0.18739203,0.046127185,0.0060753743,0.02002257,0.0067082904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012143417,0.00023830915,0.00025996548,0.0020350183,0.0013028416,0.0015893566,0.00068916776,0.000534763,0.003833749],"category_scores_gemma":[0.10567587,0.00023770619,0.00030719375,0.002181708,0.0018856765,0.0026143568,0.001994651,0.00092079386,0.0004911431],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028444978,0.00026064645,0.9520483,0.00017149455,0.000048286576,0.00017258248,0.020378005,0.00019889446,0.0004663597,0.0009591003,0.00061756145,0.024394406],"study_design_scores_gemma":[0.000023778824,0.00051327125,0.9683732,0.00017810494,0.00004171487,0.00029079252,0.024903448,0.0012178521,0.00076387904,0.0008600014,0.002809854,0.000024089337],"about_ca_topic_score_codex":0.0034369782,"about_ca_topic_score_gemma":0.0057689263,"teacher_disagreement_score":0.012143417,"about_ca_system_score_codex":0.0015503088,"about_ca_system_score_gemma":0.0019589877,"threshold_uncertainty_score":0.06422132},"labels":[],"label_agreement":null},{"id":"W4412704074","doi":"10.1145/3696630.3728518","title":"From Overload to Insight: Bridging Code Search and Code Review with LLMs","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Bridging (networking); Computer science; Information overload; Code (set theory); Programming language; Computer security; World Wide Web","score_opus":0.0195272804877584,"score_gpt":0.2974078624599623,"score_spread":0.2778805819722039,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412704074","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06481886,0.005953672,0.8290626,0.052943286,0.001937246,0.0016491171,0.00028054818,0.011708916,0.03164578],"genre_scores_gemma":[0.55404806,0.0019877697,0.41478488,0.009862228,0.0026898577,0.0021262106,0.00038866268,0.0021800245,0.011932322],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.79069847,0.17181174,0.007150197,0.009504745,0.017917389,0.0029174401],"domain_scores_gemma":[0.5143834,0.3737913,0.0341418,0.03743374,0.028755587,0.011494235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.09260549,0.0017396872,0.0018569442,0.011034028,0.00604061,0.015673423,0.005326723,0.0055657616,0.008125996],"category_scores_gemma":[0.32006142,0.0018214509,0.0010838438,0.0038458654,0.008641199,0.023282483,0.027984334,0.004533983,0.0047995173],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018415138,0.00050457724,0.010335836,0.0037606275,0.00022183693,0.0019953733,0.16430187,0.00277921,0.022156807,0.08829262,0.03519425,0.6686154],"study_design_scores_gemma":[0.000526114,0.0010778848,0.005911162,0.0031678183,0.00042572178,0.002390585,0.04555167,0.05906064,0.02014523,0.36831167,0.49274915,0.0006823674],"about_ca_topic_score_codex":0.0013533463,"about_ca_topic_score_gemma":0.0021251827,"teacher_disagreement_score":0.09260549,"about_ca_system_score_codex":0.003737321,"about_ca_system_score_gemma":0.01244582,"threshold_uncertainty_score":0.4897505},"labels":[],"label_agreement":null},{"id":"W4412704076","doi":"10.1145/3696630.3728591","title":"TS-Detector : Detecting Feature Toggle Usage Patterns","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; University of Saskatchewan","funders":"","keywords":"Computer science; Feature (linguistics); Detector; Pattern recognition (psychology); Artificial intelligence; Telecommunications","score_opus":0.020952586719881247,"score_gpt":0.284319805634565,"score_spread":0.26336721891468373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412704076","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45728898,0.0008405269,0.32894966,0.0005150698,0.00026620238,0.0006268229,0.016193382,0.18875727,0.006562041],"genre_scores_gemma":[0.7348726,0.0003052225,0.23610952,0.00032459796,0.000050313014,0.0006509383,0.015308115,0.004936581,0.0074422164],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9959699,0.0006243631,0.0004705671,0.0011108462,0.0015564684,0.00026796488],"domain_scores_gemma":[0.98136526,0.010228809,0.003457439,0.0020211097,0.0024191637,0.0005081682],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022014428,0.0016176623,0.0008573712,0.007678926,0.0005234762,0.0013607926,0.0014867187,0.0013631695,0.002768695],"category_scores_gemma":[0.017945634,0.0006899359,0.000847575,0.0024206329,0.00052255794,0.0020244035,0.0020852496,0.0009335215,0.0020831265],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012071705,0.0006059138,0.286703,0.0018175187,0.00045474686,0.0021999865,0.0033289148,0.008957837,0.0774257,0.002826,0.07769468,0.5367785],"study_design_scores_gemma":[0.00031619507,0.0012685988,0.1991359,0.0005835944,0.00035778125,0.0054359436,0.002111863,0.4791923,0.24075393,0.008775369,0.06159273,0.00047585389],"about_ca_topic_score_codex":0.0030569432,"about_ca_topic_score_gemma":0.004072019,"teacher_disagreement_score":0.007678926,"about_ca_system_score_codex":0.00046992698,"about_ca_system_score_gemma":0.00086656626,"threshold_uncertainty_score":0.011642456},"labels":[],"label_agreement":null},{"id":"W4412714478","doi":"10.59543/ijmscs.v3i.15097","title":"Comparative Analysis of AI Models for Effort Estimation in Western and Regional Environments","year":2025,"lang":"en","type":"article","venue":"International Journal of Mathematics Statistics and Computer Science","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Mean squared error; Random forest; Workflow; Estimation; Artificial neural network; Work (physics); Software; Computer science; Documentation; Operations research; Artificial intelligence; Statistics; Engineering; Mathematics; Economics; Management","score_opus":0.028405859222562984,"score_gpt":0.3425853700757589,"score_spread":0.3141795108531959,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412714478","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7930408,0.002773135,0.18489973,0.0012473229,0.00012485588,0.00024963616,0.0011013235,0.001267896,0.015295282],"genre_scores_gemma":[0.94654775,0.0009678819,0.048186418,0.00010658307,0.000037854916,0.00018061363,0.0013198754,0.0001376772,0.002515417],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.99604553,0.0022744532,0.00023732192,0.0005524179,0.0006185515,0.0002716859],"domain_scores_gemma":[0.9656492,0.027575351,0.0020238864,0.00087357964,0.0035843381,0.0002936382],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011088489,0.0010829221,0.00085236813,0.0029273068,0.0004924522,0.002179202,0.0015449971,0.00081914297,0.0016645208],"category_scores_gemma":[0.026333,0.00038961333,0.0014468692,0.0034305332,0.00062096835,0.001863298,0.0010342598,0.001153325,0.00058838417],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000892716,0.00036005955,0.08942289,0.0003600654,0.0005430156,0.00025271182,0.00074402115,0.79306585,0.0007154687,0.0086680725,0.0024767758,0.10249837],"study_design_scores_gemma":[0.000016042128,0.00017421514,0.015031295,0.000069879745,0.0000816674,0.00006007012,0.00055498723,0.9799585,0.00067156856,0.002093982,0.0012499432,0.000037910908],"about_ca_topic_score_codex":0.042017348,"about_ca_topic_score_gemma":0.025233798,"teacher_disagreement_score":0.042017348,"about_ca_system_score_codex":0.0019715643,"about_ca_system_score_gemma":0.0015306306,"threshold_uncertainty_score":0.083545566},"labels":[],"label_agreement":null},{"id":"W4412781231","doi":"10.1145/3702652.3744207","title":"How Do Novice Programmers Solve Code-Tracing Problems When ChatGPT Is Available? A Qualitative Analysis.","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Tracing; Programming language; Code (set theory); Software engineering","score_opus":0.03257780147867444,"score_gpt":0.33006478108658166,"score_spread":0.2974869796079072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412781231","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9681234,0.00034164853,0.010996308,0.0023554377,0.00003986157,0.0004576478,0.000474862,0.000080542086,0.017130291],"genre_scores_gemma":[0.9901245,0.00025463937,0.0042351354,0.0005366541,0.000005542757,0.00033753063,0.00016341275,0.000058121648,0.0042844983],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9886895,0.0077967793,0.0003177985,0.00077470345,0.0015431058,0.00087812694],"domain_scores_gemma":[0.8431065,0.13431938,0.0065483404,0.0024502121,0.010294605,0.003280836],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014692701,0.00028990867,0.00034966154,0.0013009507,0.0023772484,0.0032512888,0.001550621,0.0016370658,0.0036293764],"category_scores_gemma":[0.1057195,0.0004475155,0.00022450495,0.0013723929,0.003541807,0.004078866,0.0023901374,0.0023295004,0.0006401437],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001936245,0.00028771028,0.021645153,0.0009993833,0.000019252771,0.00047647976,0.94216335,0.00025026425,0.003388708,0.0039200196,0.0023093591,0.024346676],"study_design_scores_gemma":[0.000042407715,0.00026406124,0.025897024,0.0007128752,0.000027700033,0.00032314574,0.948469,0.0012712725,0.0032906216,0.0029529985,0.016693857,0.000054986947],"about_ca_topic_score_codex":0.006345785,"about_ca_topic_score_gemma":0.010588462,"teacher_disagreement_score":0.014692701,"about_ca_system_score_codex":0.0035889375,"about_ca_system_score_gemma":0.0049583875,"threshold_uncertainty_score":0.07770342},"labels":[],"label_agreement":null},{"id":"W4412888061","doi":"10.18653/v1/2025.findings-acl.784","title":"Enhancing LLM Agent Safety via Causal Influence Prompting","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Human–computer interaction; Risk analysis (engineering); Business","score_opus":0.009514871464068484,"score_gpt":0.27328559283823683,"score_spread":0.26377072137416835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412888061","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0123061305,0.000105894724,0.9775514,0.00025375126,0.000043092376,0.00019802786,0.00011543264,0.0068325116,0.0025937161],"genre_scores_gemma":[0.41648924,0.00017714719,0.57834023,0.0002298716,0.000048563976,0.0003323227,0.00040569503,0.0009787168,0.0029981895],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965628,0.0012496422,0.00014874959,0.0005670458,0.0012776606,0.00019406695],"domain_scores_gemma":[0.98779094,0.008159846,0.00093727105,0.0014819186,0.0012926805,0.0003372574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036856055,0.0014802705,0.00065057963,0.0016483637,0.0009022734,0.0014424511,0.001531256,0.00093480136,0.005628103],"category_scores_gemma":[0.0243005,0.00070426264,0.0012231703,0.0005799223,0.0014391663,0.0028012628,0.0038361277,0.0020800845,0.0008808732],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007545407,0.00046867057,0.009981304,0.0008474962,0.000119391436,0.0010629903,0.0030783617,0.32158536,0.03736548,0.11028774,0.007854068,0.5065946],"study_design_scores_gemma":[0.00007825371,0.00013292185,0.0004554304,0.00006947832,0.00008045397,0.0001746515,0.00021505049,0.8967539,0.032965492,0.049270887,0.019753229,0.000050313833],"about_ca_topic_score_codex":0.004605404,"about_ca_topic_score_gemma":0.0073604686,"teacher_disagreement_score":0.005628103,"about_ca_system_score_codex":0.0012287326,"about_ca_system_score_gemma":0.0035169588,"threshold_uncertainty_score":0.019491613},"labels":[],"label_agreement":null},{"id":"W4412934941","doi":"10.1007/s00236-025-00495-x","title":"Novel tree-search method for synthesizing SMT strategies","year":2025,"lang":"en","type":"article","venue":"Acta Informatica","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Alliance de recherche numérique du Canada; Georg-August-Universität Göttingen","keywords":"Computer science; Theory of computation; Tree (set theory); Search tree; Theoretical computer science; Programming language; Parallel computing; Search algorithm; Mathematics; Combinatorics","score_opus":0.030245868436263517,"score_gpt":0.3359139184326752,"score_spread":0.3056680499964117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412934941","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008179506,0.00015246168,0.98468137,0.0001257962,0.000034240456,0.0001347632,0.0001574274,0.002023056,0.0045113536],"genre_scores_gemma":[0.13054013,0.00012691836,0.8659161,0.00013864244,0.000017008668,0.00023949464,0.00045158796,0.0005543751,0.002015726],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990507,0.00027228636,0.00007103975,0.00017395407,0.00033337492,0.000098788776],"domain_scores_gemma":[0.9981981,0.0012406346,0.00011056464,0.00014804625,0.0002434512,0.000059200407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012589392,0.001033788,0.00069072674,0.001286382,0.000446145,0.001044203,0.0012978844,0.0010810832,0.011000724],"category_scores_gemma":[0.005354581,0.0005760792,0.0012262668,0.0009693312,0.000855341,0.0011193025,0.0013474636,0.0011144748,0.0020225055],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021600636,0.00014064317,0.0017946974,0.00053815974,0.000093548995,0.0002778265,0.00022962881,0.5115232,0.012143508,0.0800471,0.008295622,0.3847],"study_design_scores_gemma":[0.000047031204,0.00004047451,0.000060473063,0.00002915798,0.00001542858,0.000042890206,0.000022951628,0.9787724,0.0026698369,0.014468586,0.003824124,0.000006631258],"about_ca_topic_score_codex":0.0035513802,"about_ca_topic_score_gemma":0.006956123,"teacher_disagreement_score":0.011000724,"about_ca_system_score_codex":0.0010494286,"about_ca_system_score_gemma":0.0027549788,"threshold_uncertainty_score":0.0368011},"labels":[],"label_agreement":null},{"id":"W4413002204","doi":"10.1016/j.jlp.2025.105751","title":"Intelligent countermeasures analysis in oil and gas projects utilizing topic modeling","year":2025,"lang":"en","type":"article","venue":"Journal of Loss Prevention in the Process Industries","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Mitacs; International Dragonfly Fund","keywords":"Petroleum engineering; Engineering; Systems engineering","score_opus":0.05576815941098868,"score_gpt":0.34506841590168935,"score_spread":0.2893002564907007,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413002204","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43625167,0.0027159408,0.5527084,0.0010431058,0.00012472185,0.00038825982,0.0020314336,0.0014090305,0.0033274344],"genre_scores_gemma":[0.8909866,0.0010679644,0.10256727,0.000076662785,0.00010173385,0.0002838365,0.0035544094,0.000062755185,0.0012986278],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983461,0.0006742924,0.00016694418,0.00039864847,0.00027360665,0.00014040653],"domain_scores_gemma":[0.99666494,0.0022853361,0.00034990828,0.00016486314,0.00045688768,0.000078149176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00263147,0.0007426706,0.00061153114,0.0039578127,0.0005332149,0.0016851954,0.00088088826,0.00077534514,0.00069675874],"category_scores_gemma":[0.006331531,0.00028221993,0.0017465942,0.0022816858,0.00030529857,0.0013052822,0.00091053394,0.0008301356,0.00029105577],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066210487,0.0005808677,0.13271156,0.0009225098,0.0007256686,0.0008935731,0.0029454003,0.19655506,0.012442443,0.011664975,0.007949984,0.63194585],"study_design_scores_gemma":[0.00002654901,0.00013785069,0.028747244,0.00009593885,0.0002732207,0.0001957031,0.0013743334,0.94985056,0.0044734497,0.008178978,0.006601728,0.000044477172],"about_ca_topic_score_codex":0.009578443,"about_ca_topic_score_gemma":0.009966104,"teacher_disagreement_score":0.009578443,"about_ca_system_score_codex":0.0008108782,"about_ca_system_score_gemma":0.0013295853,"threshold_uncertainty_score":0.019045353},"labels":[],"label_agreement":null},{"id":"W4413133129","doi":"10.36227/techrxiv.175459856.69720288/v1","title":"Multifaceted LLMs for Recommendations of Software Vulnerability Repair and Assessment","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Context (archaeology); Vulnerability (computing); Codebase; Code (set theory); Computer security; Data science; Set (abstract data type); Source code; Programming language; Geography","score_opus":0.049026545279198176,"score_gpt":0.37858464291018906,"score_spread":0.32955809763099086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413133129","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2137445,0.0056543853,0.72208303,0.003403499,0.00025489437,0.0027184326,0.013529098,0.024176868,0.014435275],"genre_scores_gemma":[0.38879883,0.00061914127,0.5987119,0.0003637334,0.00007159016,0.00078438997,0.007665517,0.000345944,0.0026388918],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9917494,0.0022668578,0.0010262548,0.0013698159,0.0032210376,0.00036673926],"domain_scores_gemma":[0.965643,0.019768294,0.0024609359,0.0036566898,0.007676982,0.0007941892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060047084,0.002012423,0.0012208818,0.011936334,0.0010576093,0.0031606697,0.0019712453,0.0024564,0.0043414026],"category_scores_gemma":[0.056992646,0.00070591026,0.0017491017,0.0044620233,0.0004922431,0.004782758,0.0019628857,0.0025346314,0.0025925844],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091790885,0.0009270444,0.065268695,0.002394907,0.00071209914,0.00043926248,0.0026629928,0.034882344,0.01460826,0.006203391,0.025134997,0.8458481],"study_design_scores_gemma":[0.00025610434,0.0012648129,0.04077258,0.0013060436,0.00086918287,0.0005794044,0.003024156,0.8475134,0.02395893,0.02209719,0.057986267,0.00037184896],"about_ca_topic_score_codex":0.018855704,"about_ca_topic_score_gemma":0.048650246,"teacher_disagreement_score":0.018855704,"about_ca_system_score_codex":0.0019546016,"about_ca_system_score_gemma":0.0029211657,"threshold_uncertainty_score":0.037491858},"labels":[],"label_agreement":null},{"id":"W4413158318","doi":"10.1109/icosse65712.2025.00017","title":"How Software Development Professionals Perceive the Use of Code Reviewer Recommendation Systems","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Software engineering; Code (set theory); Software development; Code review; Software; Programming language; Static program analysis","score_opus":0.07708941263140973,"score_gpt":0.3172868489674504,"score_spread":0.24019743633604068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413158318","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9820015,0.0009801418,0.00528882,0.0038825332,0.000084003994,0.000080547674,0.000052225765,0.00018215775,0.0074480325],"genre_scores_gemma":[0.98986065,0.0007955668,0.0063022687,0.0012177449,0.00005600873,0.000051076055,0.00010051589,0.00006959466,0.0015463445],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9563112,0.025293246,0.0027877497,0.0020170591,0.012016222,0.0015744973],"domain_scores_gemma":[0.70320904,0.19206522,0.030915312,0.0077668997,0.04953306,0.016510446],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.033198375,0.0004290017,0.00041428648,0.003595459,0.0017256144,0.0056575984,0.0009790238,0.0023734188,0.0018998059],"category_scores_gemma":[0.2005698,0.00062009366,0.00047969853,0.0013837531,0.00168351,0.0037521932,0.0023096928,0.0015278239,0.0010427922],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004132061,0.00058858225,0.5548074,0.000750863,0.00020992284,0.0008591893,0.23017606,0.00050119555,0.005917166,0.0006775771,0.009015625,0.19608319],"study_design_scores_gemma":[0.0001887873,0.002785962,0.52244025,0.0014407514,0.00027046786,0.003883501,0.40524933,0.007834634,0.0029286372,0.001973643,0.0504857,0.0005183352],"about_ca_topic_score_codex":0.006153488,"about_ca_topic_score_gemma":0.0069794566,"teacher_disagreement_score":0.96680164,"about_ca_system_score_codex":0.001758854,"about_ca_system_score_gemma":0.003238283,"threshold_uncertainty_score":0.17557186},"labels":[],"label_agreement":null},{"id":"W4413216650","doi":"10.1145/3712255.3726730","title":"Learning to Predict Code Review Rounds in Modern Code Review Using Multi-Objective Genetic Programming","year":2025,"lang":"en","type":"article","venue":"Proceedings of the Genetic and Evolutionary Computation Conference Companion","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Concordia University","funders":"","keywords":"Computer science; Genetic programming; Programming language; Code (set theory); Code review; Artificial intelligence; Static program analysis; Software development; Software","score_opus":0.034979321908794485,"score_gpt":0.3047175831055146,"score_spread":0.2697382611967201,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413216650","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3529947,0.0015531001,0.635344,0.0021665525,0.00013150265,0.00043061798,0.0010266582,0.0032655224,0.0030872694],"genre_scores_gemma":[0.84806436,0.00029451956,0.14635904,0.00048070698,0.0000838096,0.0003130158,0.0017058438,0.00021396573,0.0024847374],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99786645,0.0008246903,0.00012247212,0.0006332534,0.00036009774,0.00019306452],"domain_scores_gemma":[0.97797775,0.016322874,0.0021787027,0.0005657162,0.0024416817,0.0005132541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005148545,0.0012603645,0.0012073122,0.0027035319,0.0005454607,0.0017764972,0.0016545289,0.0014626321,0.0012531171],"category_scores_gemma":[0.02072107,0.00074627926,0.0008785337,0.0015381727,0.00061790604,0.0015845916,0.0009261024,0.0019286778,0.000519244],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018097073,0.00024176131,0.043075584,0.00021667128,0.00016915913,0.00016951798,0.000258783,0.77213967,0.0013546463,0.0021027976,0.0043547573,0.17573568],"study_design_scores_gemma":[0.000008055419,0.000028163944,0.0016194902,0.000013393172,0.000014308826,0.000015683416,0.000022424936,0.9956233,0.0004461205,0.0019004692,0.00030125974,0.000007350522],"about_ca_topic_score_codex":0.013389289,"about_ca_topic_score_gemma":0.021257075,"teacher_disagreement_score":0.013389289,"about_ca_system_score_codex":0.0019828207,"about_ca_system_score_gemma":0.002652995,"threshold_uncertainty_score":0.027228415},"labels":[],"label_agreement":null},{"id":"W4413268009","doi":"10.1145/3696630.3731666","title":"Get on the Train or be Left on the Station: Using LLMs for Software Engineering Research","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software; Computer science; Engineering management; Engineering; Operating system","score_opus":0.1492130351029197,"score_gpt":0.38416457518612823,"score_spread":0.23495154008320854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413268009","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05158018,0.007852958,0.58385235,0.19059727,0.001501188,0.00034020987,0.00014668051,0.0012264053,0.16290288],"genre_scores_gemma":[0.73295915,0.0053578652,0.23747416,0.007789049,0.000691915,0.0009632581,0.00011594015,0.0006071592,0.014041473],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9337637,0.056559708,0.0017544588,0.0018173936,0.0051439344,0.00096092775],"domain_scores_gemma":[0.763101,0.20075783,0.007622266,0.019599538,0.006540195,0.0023792237],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.058241356,0.0006023458,0.00049020443,0.004454658,0.0058048293,0.022786353,0.0027867998,0.004668321,0.004708551],"category_scores_gemma":[0.11811228,0.00075738673,0.0008820121,0.0033867578,0.048514213,0.042271715,0.012471922,0.0054562567,0.0013057528],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026077296,0.000023254312,0.0007217143,0.00014950892,0.0000070367446,0.0001312,0.013747065,0.0009579346,0.00019742596,0.94547063,0.0025123341,0.03605582],"study_design_scores_gemma":[0.000017020684,0.000048068254,0.000373913,0.0007343657,0.000018314207,0.00015670972,0.0097822,0.004774723,0.0009447787,0.9030731,0.080030136,0.00004657611],"about_ca_topic_score_codex":0.004909173,"about_ca_topic_score_gemma":0.0052380147,"teacher_disagreement_score":0.94175863,"about_ca_system_score_codex":0.00932908,"about_ca_system_score_gemma":0.010923291,"threshold_uncertainty_score":0.30801338},"labels":[],"label_agreement":null},{"id":"W4413279387","doi":"10.1145/3760775","title":"VulScribeR: Exploring RAG-based Vulnerability Augmentation with LLMs","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Vulnerability (computing); Computer security","score_opus":0.12021106659626543,"score_gpt":0.33913627368043525,"score_spread":0.21892520708416982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413279387","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25230625,0.0050842455,0.6538383,0.0021713893,0.00048864575,0.00055261457,0.0049738633,0.07451726,0.00606738],"genre_scores_gemma":[0.6058777,0.00079185975,0.37088004,0.001686972,0.00015660973,0.0005356091,0.013440086,0.0017764384,0.0048546633],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984824,0.00043618606,0.00008876874,0.00051470497,0.00034966107,0.00012826207],"domain_scores_gemma":[0.9964347,0.0018772655,0.0002432755,0.0010003949,0.0003359636,0.000108399036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020531896,0.001909813,0.0012141464,0.0025431847,0.0004973365,0.0010530425,0.0023124963,0.0015883883,0.0021620276],"category_scores_gemma":[0.008109805,0.0005541764,0.0017533213,0.0012997227,0.0012539583,0.003550238,0.0027323114,0.0020400789,0.0014466564],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005849528,0.0006304093,0.015779292,0.00062634744,0.00027361268,0.0005864232,0.00051753933,0.23441571,0.026069375,0.005791798,0.03391153,0.68081295],"study_design_scores_gemma":[0.00006953746,0.00027496277,0.0014265575,0.000052429394,0.00009398221,0.0003332731,0.00012972613,0.9604155,0.014662179,0.012625504,0.009872499,0.000043829925],"about_ca_topic_score_codex":0.0029385933,"about_ca_topic_score_gemma":0.0054358607,"teacher_disagreement_score":0.0029385933,"about_ca_system_score_codex":0.0009055794,"about_ca_system_score_gemma":0.0014362774,"threshold_uncertainty_score":0.010858417},"labels":[],"label_agreement":null},{"id":"W4413308446","doi":"10.1007/s10664-025-10707-0","title":"Towards understanding the challenges of bug localization in deep learning systems","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Deep learning; Data science; Artificial intelligence; Software engineering","score_opus":0.042013341184440375,"score_gpt":0.2910416336842223,"score_spread":0.24902829249978192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413308446","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11640395,0.0024975212,0.86499465,0.011404836,0.00012660715,0.00007291309,0.00033422455,0.000788277,0.003377018],"genre_scores_gemma":[0.87095386,0.0011758065,0.124076515,0.0007706493,0.0001926187,0.0000641817,0.0003710143,0.00018566758,0.0022095712],"study_design_codex":"simulation_or_modeling","study_design_gemma":"qualitative","domain_scores_codex":[0.9964162,0.0015419251,0.00022376298,0.0006759693,0.0007522824,0.00038987925],"domain_scores_gemma":[0.9533638,0.03423564,0.0037107451,0.003877431,0.0039573465,0.0008549324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067130523,0.00090145227,0.0011963465,0.0017783119,0.0008003928,0.0050304276,0.0021304039,0.0029662368,0.0024875957],"category_scores_gemma":[0.06177506,0.0009890986,0.0006377535,0.0012431311,0.0026997481,0.01264092,0.0041648815,0.00580579,0.0004492895],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003445415,0.00045262117,0.04120562,0.000695244,0.00023250844,0.00028813956,0.0016888736,0.41670752,0.004926101,0.23345998,0.0085484125,0.29145044],"study_design_scores_gemma":[0.000008369932,0.000025649511,0.0011199042,0.000058703525,0.00001414218,0.000035258327,0.00020193304,0.77353334,0.0008193106,0.22290859,0.0012627535,0.000011988563],"about_ca_topic_score_codex":0.008028561,"about_ca_topic_score_gemma":0.008116927,"teacher_disagreement_score":0.008028561,"about_ca_system_score_codex":0.0018559756,"about_ca_system_score_gemma":0.002961522,"threshold_uncertainty_score":0.035502434},"labels":[],"label_agreement":null},{"id":"W4413333491","doi":"10.1145/3762183","title":"Leveraging Reviewer Experience in Code Review Comment Generation","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Code (set theory); Software engineering; Programming language","score_opus":0.17826252847388172,"score_gpt":0.3910005918240723,"score_spread":0.2127380633501906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413333491","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41551328,0.0029520395,0.5561791,0.0031886115,0.0010243222,0.0016817262,0.0009371546,0.010784139,0.0077395225],"genre_scores_gemma":[0.90706277,0.00028798854,0.08696821,0.0006413994,0.00034448426,0.000498991,0.0010039367,0.00045287266,0.0027393322],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96440035,0.020896586,0.0025708082,0.004998018,0.0064085517,0.00072563434],"domain_scores_gemma":[0.7421616,0.17123803,0.02003005,0.01729522,0.045766104,0.003509014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03487788,0.001804577,0.0014370716,0.0037426387,0.0007179347,0.00291421,0.00180178,0.0023518775,0.0016101871],"category_scores_gemma":[0.20015691,0.0007194034,0.0009334161,0.0013189581,0.00085994345,0.0039735264,0.0023457075,0.0020724514,0.0016160565],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023605304,0.001009451,0.11234104,0.001630662,0.00061795895,0.0008017747,0.0040974882,0.071349725,0.028376993,0.0017918912,0.017685965,0.75793654],"study_design_scores_gemma":[0.00035640973,0.0016807346,0.030544937,0.00037229332,0.00037431295,0.0008566565,0.00086443854,0.9157124,0.030987099,0.0058248863,0.012094614,0.0003311419],"about_ca_topic_score_codex":0.001424646,"about_ca_topic_score_gemma":0.0025371788,"teacher_disagreement_score":0.03487788,"about_ca_system_score_codex":0.0010875621,"about_ca_system_score_gemma":0.0016837528,"threshold_uncertainty_score":0.18445408},"labels":[],"label_agreement":null},{"id":"W4413639962","doi":"10.1109/compsac65507.2025.00172","title":"A Contrastive Learning Approach to Bug Severity Classification with Large Language Model Embeddings","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Natural language processing","score_opus":0.013382062783903047,"score_gpt":0.2759558994778363,"score_spread":0.26257383669393325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413639962","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1992593,0.0014866764,0.7853906,0.0008555085,0.00023896073,0.00022748123,0.0019011237,0.008858842,0.0017813867],"genre_scores_gemma":[0.67087865,0.0003342431,0.31918108,0.00043603344,0.00019428089,0.0002451326,0.005962816,0.00037354385,0.0023942592],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988557,0.00040624803,0.00009198545,0.0003786937,0.00019430337,0.00007298032],"domain_scores_gemma":[0.9964139,0.0017862329,0.00039466206,0.0006363248,0.0006420144,0.00012680168],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013411172,0.0014067647,0.0007383476,0.002173254,0.00037819095,0.00092939334,0.0011222,0.0009124794,0.0010065554],"category_scores_gemma":[0.007633514,0.00037520583,0.00094561046,0.0015205808,0.0005030049,0.0027119427,0.0017102144,0.0020833425,0.0009883713],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060242147,0.0013193146,0.021111757,0.00032105882,0.00033424175,0.00031839905,0.0004734531,0.09782527,0.02461473,0.0044018147,0.015463384,0.83321416],"study_design_scores_gemma":[0.000052149186,0.0003178825,0.0022644075,0.000020925636,0.000052012103,0.00013955022,0.000106516716,0.980354,0.005859046,0.00833021,0.002469444,0.000033836273],"about_ca_topic_score_codex":0.002143388,"about_ca_topic_score_gemma":0.004828445,"teacher_disagreement_score":0.002173254,"about_ca_system_score_codex":0.00063428294,"about_ca_system_score_gemma":0.00076993764,"threshold_uncertainty_score":0.007092595},"labels":[],"label_agreement":null},{"id":"W4413679385","doi":"10.1109/compsac65507.2025.00165","title":"Predicting Bug Inducing Commits Using Commit-State Transition Analysis","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Western University","funders":"","keywords":"Commit; Computer science; Transition (genetics); State (computer science); Programming language; Database; Chemistry","score_opus":0.020298312882595333,"score_gpt":0.28871732484942364,"score_spread":0.2684190119668283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413679385","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7805896,0.0025850707,0.14121889,0.0007575468,0.0002733209,0.0004198361,0.04297225,0.02809751,0.0030860312],"genre_scores_gemma":[0.7724303,0.0005054566,0.104797006,0.00011386876,0.00007394596,0.0002683349,0.118823245,0.0007235272,0.002264254],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986327,0.00015890734,0.00014019165,0.00047317074,0.00048056574,0.000114353454],"domain_scores_gemma":[0.989051,0.004412005,0.0017728105,0.0014739407,0.0026985279,0.00059180934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018763436,0.0011357194,0.0005825282,0.005594512,0.000568273,0.00097108184,0.0012289528,0.0009511113,0.0007095527],"category_scores_gemma":[0.011326049,0.00041882522,0.0008289496,0.0026014848,0.0003964799,0.0011713086,0.0011034197,0.0014263944,0.00074711646],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092439615,0.0007905799,0.46081924,0.0010925684,0.0003552165,0.0013275503,0.00074115244,0.13388401,0.018301161,0.004671159,0.064713985,0.31237897],"study_design_scores_gemma":[0.00010238936,0.00039907353,0.100429274,0.00008247767,0.00012861806,0.00058739283,0.00030545623,0.8634796,0.012903684,0.006762543,0.014744933,0.000074502204],"about_ca_topic_score_codex":0.014041257,"about_ca_topic_score_gemma":0.028879657,"teacher_disagreement_score":0.014041257,"about_ca_system_score_codex":0.00074552547,"about_ca_system_score_gemma":0.0015678388,"threshold_uncertainty_score":0.027919054},"labels":[],"label_agreement":null},{"id":"W4413679657","doi":"10.1109/compsac65507.2025.00155","title":"Program Slicing in the Era of Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Slicing; Computer science; Program slicing; Programming language; Natural language processing; World Wide Web","score_opus":0.013569981210575432,"score_gpt":0.31692709827042065,"score_spread":0.3033571170598452,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413679657","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19097823,0.005872655,0.7518896,0.0042167404,0.00024434883,0.0002782376,0.0026479675,0.04204413,0.0018280919],"genre_scores_gemma":[0.50264215,0.0014273329,0.48544058,0.0012242517,0.00013885369,0.00033699465,0.0056371787,0.0021304537,0.0010221653],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9941724,0.0026529771,0.00030852715,0.0012612606,0.0013833867,0.00022146339],"domain_scores_gemma":[0.9599303,0.02891946,0.0026646792,0.006108948,0.0017988266,0.00057778787],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008600446,0.0015709425,0.0011497069,0.0017572623,0.0007587329,0.0023016755,0.0026788472,0.0013238342,0.0012309572],"category_scores_gemma":[0.053128917,0.0013121708,0.0016743643,0.0014693177,0.0018885866,0.0074701584,0.0036005087,0.0037281322,0.00052301784],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014682882,0.00040309824,0.03957918,0.002201269,0.00055574847,0.0015653154,0.004279701,0.40694106,0.022171162,0.03063591,0.027709499,0.4624898],"study_design_scores_gemma":[0.000098878096,0.00014900873,0.0018701615,0.00012864468,0.000084350264,0.00027955076,0.00030602203,0.9297026,0.007204348,0.049652033,0.010469391,0.00005503253],"about_ca_topic_score_codex":0.010631334,"about_ca_topic_score_gemma":0.020726886,"teacher_disagreement_score":0.010631334,"about_ca_system_score_codex":0.0017666657,"about_ca_system_score_gemma":0.003235565,"threshold_uncertainty_score":0.045484066},"labels":[],"label_agreement":null},{"id":"W4414359425","doi":"10.24963/ijcai.2025/1186","title":"Large Language Models for Causal Discovery: Current Landscape and Future Directions","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Transformative learning; Causality (physics); Leverage (statistics); Key (lock); Natural language; Metadata; Causal model","score_opus":0.010765958412345466,"score_gpt":0.2922617763234516,"score_spread":0.28149581791110617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414359425","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005308195,0.37002814,0.45235622,0.15739939,0.0007921136,0.00022287866,0.0011419725,0.001795031,0.010955994],"genre_scores_gemma":[0.12956792,0.37533516,0.4672419,0.014588207,0.0067193694,0.0013294141,0.002015702,0.0006459708,0.002556347],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.97296596,0.020991769,0.0008697525,0.002212659,0.0025558427,0.00040399237],"domain_scores_gemma":[0.526693,0.44997963,0.003293498,0.010137904,0.0080877775,0.0018081327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06840863,0.0017436235,0.003960904,0.006560255,0.0017261566,0.011568691,0.006267826,0.005524188,0.0118650505],"category_scores_gemma":[0.16354586,0.0020923037,0.0030644522,0.007625748,0.009044569,0.029485082,0.0070969495,0.01199919,0.0031221048],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000247021,0.00022329685,0.00549556,0.004539087,0.00043405785,0.00013486463,0.0010012692,0.022357404,0.00020188125,0.5375536,0.021332912,0.406479],"study_design_scores_gemma":[0.000054001648,0.00006055856,0.0006550028,0.0021822904,0.000099856756,0.0000982209,0.0006357087,0.06290155,0.00014203381,0.8788,0.054293133,0.000077703946],"about_ca_topic_score_codex":0.011346315,"about_ca_topic_score_gemma":0.009821074,"teacher_disagreement_score":0.06840863,"about_ca_system_score_codex":0.0061533856,"about_ca_system_score_gemma":0.010397559,"threshold_uncertainty_score":0.36178374},"labels":[],"label_agreement":null},{"id":"W4414360800","doi":"10.24963/ijcai.2024/1186","title":"Large Language Models for Causal Discovery: Current Landscape and Future Directions","year":2024,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Transformative learning; Causality (physics); Leverage (statistics); Key (lock); Natural language; Metadata; Causal model","score_opus":0.013572489784373313,"score_gpt":0.2941172482840146,"score_spread":0.2805447584996413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414360800","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053169574,0.37637728,0.4449205,0.15845747,0.00079089,0.00022141801,0.0011290692,0.001782679,0.011003672],"genre_scores_gemma":[0.1294744,0.38064253,0.4620562,0.014582428,0.006735372,0.0013200386,0.0019935116,0.0006414252,0.0025540667],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9732213,0.020774493,0.00086288416,0.002205238,0.0025317003,0.00040426993],"domain_scores_gemma":[0.52828765,0.44845247,0.0032787484,0.010084218,0.008085868,0.0018110301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.068409935,0.0017400607,0.0039405976,0.0065421155,0.0017242858,0.01152709,0.0062661157,0.00554372,0.011915684],"category_scores_gemma":[0.16246371,0.0020859712,0.0030481508,0.0075830524,0.009069862,0.029502664,0.0070608226,0.011990437,0.0031307435],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024746172,0.00022421179,0.005490927,0.0045472896,0.00042947647,0.00013418836,0.000998297,0.022103216,0.00020187379,0.5356007,0.02129685,0.40872547],"study_design_scores_gemma":[0.000054243526,0.00006142272,0.0006610388,0.0022120308,0.00010010679,0.00009869877,0.0006454408,0.062334847,0.00014298913,0.8786862,0.054925133,0.00007800407],"about_ca_topic_score_codex":0.011328901,"about_ca_topic_score_gemma":0.009783834,"teacher_disagreement_score":0.068409935,"about_ca_system_score_codex":0.006141488,"about_ca_system_score_gemma":0.01040093,"threshold_uncertainty_score":0.36179066},"labels":[],"label_agreement":null},{"id":"W4414589104","doi":"10.1007/s10664-025-10717-y","title":"Towards understanding the impact of data bugs on deep learning models in software engineering","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Dalhousie University","funders":"","keywords":"Deep learning; Overfitting; Generalizability theory; Data pre-processing; Software quality; Leverage (statistics); Preprocessor; Metric (unit); Software bug; Software","score_opus":0.09277203969085854,"score_gpt":0.34203059390168095,"score_spread":0.2492585542108224,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414589104","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44974697,0.0028287126,0.5333462,0.008423022,0.0002009618,0.00005858683,0.0004070438,0.001048957,0.0039395452],"genre_scores_gemma":[0.9653669,0.0005115622,0.03231659,0.00031200782,0.00007475136,0.000026768892,0.0002349541,0.00014225715,0.0010141599],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.996784,0.0016764044,0.00018395834,0.00044289522,0.0006351525,0.00027763774],"domain_scores_gemma":[0.8846455,0.09872138,0.0051708003,0.0051543806,0.0053420872,0.0009657761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009754561,0.0009840745,0.0009062155,0.0013044184,0.0005024909,0.0031018576,0.0016104279,0.001975525,0.0024209816],"category_scores_gemma":[0.12418007,0.00093298923,0.0005857886,0.0011112054,0.0016055813,0.010612553,0.0024051664,0.005790864,0.0002627426],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043611124,0.00037596142,0.035379935,0.0002871477,0.00018765785,0.00012997551,0.00038251217,0.7907699,0.0022905818,0.05901787,0.0030780686,0.10766425],"study_design_scores_gemma":[0.000005153065,0.000018827875,0.0006075088,0.000019652343,0.000010890816,0.0000067200626,0.000022863087,0.9741416,0.00039384872,0.024647178,0.00012125744,0.0000045034863],"about_ca_topic_score_codex":0.011293578,"about_ca_topic_score_gemma":0.011947529,"teacher_disagreement_score":0.011293578,"about_ca_system_score_codex":0.0024504869,"about_ca_system_score_gemma":0.0021078335,"threshold_uncertainty_score":0.05158764},"labels":[],"label_agreement":null},{"id":"W4414608058","doi":"10.1016/j.jss.2025.112636","title":"BugMentor: Generating answers to follow-up questions from software bug reports using structured information retrieval and neural text generation","year":2025,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Question answering; Context (archaeology); Similarity (geometry); Software; Semantics (computer science); Language model; Software bug; Face (sociological concept)","score_opus":0.017431393734662102,"score_gpt":0.26491104133831883,"score_spread":0.24747964760365673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414608058","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09361492,0.0019935758,0.78958493,0.0014347424,0.00046410834,0.0016875716,0.0042468556,0.09990604,0.0070671695],"genre_scores_gemma":[0.28856498,0.00052071764,0.688886,0.0006420761,0.00014751584,0.00095895963,0.010009838,0.0015184599,0.008751542],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99779,0.00084095314,0.00016390794,0.00052616646,0.0005845318,0.00009444118],"domain_scores_gemma":[0.9940382,0.0036284348,0.0005510179,0.0007155647,0.00091929146,0.00014746278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001981899,0.0021916546,0.000838169,0.0023787618,0.00042058586,0.0009301375,0.0018334849,0.0019181874,0.0068789185],"category_scores_gemma":[0.013090206,0.00040454452,0.0011534002,0.0009005217,0.00050578575,0.0017805108,0.0012864555,0.0012916181,0.002860746],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046601373,0.0007741276,0.0034418912,0.0015457332,0.00014120301,0.00067922974,0.0011412352,0.022323102,0.03884212,0.0038194887,0.04082285,0.8860031],"study_design_scores_gemma":[0.0005908372,0.0013005332,0.004818237,0.00022723454,0.00022049903,0.0011728505,0.0006770553,0.849291,0.082426906,0.014866146,0.04421441,0.0001942067],"about_ca_topic_score_codex":0.0026482039,"about_ca_topic_score_gemma":0.00397313,"teacher_disagreement_score":0.0068789185,"about_ca_system_score_codex":0.000764516,"about_ca_system_score_gemma":0.0010945309,"threshold_uncertainty_score":0.02301228},"labels":[],"label_agreement":null},{"id":"W4414622738","doi":"10.1007/s10664-025-10731-0","title":"DeepCodeProbe: Evaluating Code Representation Quality in Models Trained on Code","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Polytechnique Montréal","funders":"Fonds de Recherche du Québec - Santé; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Representation (politics); Software quality; Code (set theory); Quality (philosophy); Code smell; Source code; Software; Static program analysis","score_opus":0.11928940518778033,"score_gpt":0.4194135589370064,"score_spread":0.3001241537492261,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414622738","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.72486836,0.0034139727,0.22329989,0.0015063958,0.0007809361,0.00029167676,0.0059097735,0.03560763,0.0043213083],"genre_scores_gemma":[0.8725873,0.0005239064,0.09947149,0.000653642,0.00008268717,0.00019507603,0.020913005,0.001876179,0.0036966733],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99771225,0.00077950605,0.00015064441,0.0006595314,0.00049807003,0.00019987428],"domain_scores_gemma":[0.98509395,0.009859195,0.00058900216,0.002113223,0.0019538936,0.00039075656],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004536868,0.0021885647,0.00084718026,0.0016844254,0.00054698193,0.0014158101,0.0025532236,0.0027907172,0.0031020893],"category_scores_gemma":[0.021560043,0.00071277644,0.0012222892,0.0010509909,0.0011018772,0.0035465437,0.0018913839,0.003021563,0.0014038968],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022729915,0.001148936,0.031169748,0.0007853502,0.0008048857,0.0002521475,0.00028333112,0.58698905,0.015808167,0.0028098885,0.035948817,0.32172662],"study_design_scores_gemma":[0.00010078915,0.00027042243,0.0015020577,0.000043560558,0.000058417572,0.000049204165,0.00005531287,0.98896307,0.0061394684,0.0018117634,0.0009886229,0.000017368122],"about_ca_topic_score_codex":0.013716951,"about_ca_topic_score_gemma":0.018611029,"teacher_disagreement_score":0.99546313,"about_ca_system_score_codex":0.0017721774,"about_ca_system_score_gemma":0.0024246639,"threshold_uncertainty_score":0.027274191},"labels":[],"label_agreement":null},{"id":"W4414814497","doi":"10.22214/ijraset.2025.74423","title":"Evaluating Multi-Agent AI Systems for Automated Bug Detection and Code Refactoring","year":2025,"lang":"en","type":"article","venue":"International Journal for Research in Applied Science and Engineering Technology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Code refactoring; Dataflow; Pipeline (software); Scalability; Code (set theory); Debugging; False positive paradox; Software; Abstraction","score_opus":0.10929675952370653,"score_gpt":0.46700507071813874,"score_spread":0.3577083111944322,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414814497","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7567882,0.0011838381,0.2180896,0.0005778711,0.00022551905,0.0011679374,0.00056811306,0.012645932,0.008753076],"genre_scores_gemma":[0.83803666,0.00017431335,0.15909761,0.00010434605,0.000026376627,0.0002903666,0.00063317886,0.00021191961,0.0014251879],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99339855,0.002699917,0.0005292153,0.0011636348,0.0018346867,0.00037393783],"domain_scores_gemma":[0.9789336,0.011441649,0.001839886,0.003250047,0.0034273039,0.001107614],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00841913,0.0012394565,0.0006318571,0.0015501287,0.00067401223,0.0016468532,0.0025495018,0.0013810727,0.0021473854],"category_scores_gemma":[0.021572104,0.00052544585,0.0005857021,0.0007463831,0.000882547,0.0017781716,0.0022079002,0.001099253,0.0005769947],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00487267,0.003273121,0.035744064,0.0016956705,0.0009247828,0.00031803042,0.0008863693,0.58527565,0.043713197,0.0073596877,0.004489544,0.31144723],"study_design_scores_gemma":[0.00023897165,0.001271295,0.00399914,0.000038742586,0.000121465244,0.000061261366,0.00013880566,0.976825,0.013271704,0.0018251343,0.0021735888,0.000034852823],"about_ca_topic_score_codex":0.008176761,"about_ca_topic_score_gemma":0.0063851583,"teacher_disagreement_score":0.00841913,"about_ca_system_score_codex":0.0019941279,"about_ca_system_score_gemma":0.0029478888,"threshold_uncertainty_score":0.044525146},"labels":[],"label_agreement":null},{"id":"W4415004451","doi":"10.1109/re63999.2025.00048","title":"Towards Extracting Software Requirements from App Reviews using Seq2seq Framework","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Alberta Innovates","keywords":"Task (project management); Context (archaeology); Set (abstract data type); Software; Encoder; Key (lock)","score_opus":0.07716225890123486,"score_gpt":0.3710144746411982,"score_spread":0.29385221573996334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415004451","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043106552,0.0014103734,0.9266262,0.0008027961,0.00030100002,0.00066434033,0.0110482415,0.0131886825,0.0028518012],"genre_scores_gemma":[0.13896528,0.0007684623,0.80190945,0.0007832692,0.00023904543,0.000785912,0.048773978,0.00089675596,0.006877956],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99713933,0.0010156666,0.0002688514,0.00083347555,0.00062344025,0.000119152806],"domain_scores_gemma":[0.993594,0.002614504,0.00055129133,0.00063237187,0.0024133937,0.00019445646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020004362,0.0020161404,0.0009344056,0.0030679784,0.00044689095,0.00088456465,0.0012049834,0.0012130932,0.0018908303],"category_scores_gemma":[0.0073890886,0.00068198435,0.0014796935,0.0015206068,0.0003873415,0.0018636776,0.001398244,0.0014944759,0.0040265014],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005705714,0.00037521613,0.015410616,0.0023055468,0.00027132253,0.00231493,0.0017526465,0.029228037,0.094004914,0.005115279,0.08190099,0.7667499],"study_design_scores_gemma":[0.00008997111,0.0005591273,0.01527605,0.00019714737,0.00024581174,0.0021026877,0.0009483601,0.80959,0.06883893,0.013109584,0.08886356,0.00017875188],"about_ca_topic_score_codex":0.006517434,"about_ca_topic_score_gemma":0.015947767,"teacher_disagreement_score":0.006517434,"about_ca_system_score_codex":0.00065830234,"about_ca_system_score_gemma":0.00239184,"threshold_uncertainty_score":0.012959003},"labels":[],"label_agreement":null},{"id":"W4415034806","doi":"","title":"LIS at CheckThat! 2025: Multi-Stage Open-Source Large Language Models for Fact-Checking Numerical Claims","year":2025,"lang":"en","type":"article","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Natural language; Field (mathematics); Identification (biology); Narrative; Government (linguistics); Context (archaeology)","score_opus":0.034971818971176766,"score_gpt":0.30400439040225774,"score_spread":0.26903257143108095,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415034806","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010295817,0.00012102357,0.72132087,0.0014554574,0.0004269213,0.0003458559,0.0099652065,0.24306427,0.013004661],"genre_scores_gemma":[0.3532453,0.00027290566,0.514173,0.0010319566,0.00024233747,0.00073922216,0.026606545,0.08257252,0.021116147],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.994724,0.0016177905,0.00041150095,0.0009902837,0.0017498392,0.0005065117],"domain_scores_gemma":[0.97639227,0.013245706,0.0009134552,0.0063763005,0.002524743,0.00054748985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005301083,0.0019443294,0.0011227883,0.0020766053,0.0016630992,0.005874323,0.004627073,0.0036212397,0.06379175],"category_scores_gemma":[0.03362006,0.0028130943,0.00510824,0.0010604762,0.0027997128,0.013294013,0.0072370386,0.005079635,0.019714538],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024968404,0.0006572949,0.010179864,0.0020722623,0.00042639507,0.0014321763,0.0022207876,0.12006886,0.014289654,0.4541016,0.2648135,0.1272407],"study_design_scores_gemma":[0.00032360273,0.00015743173,0.0008449902,0.00032816047,0.00014650408,0.000293249,0.00028307392,0.6524198,0.023070185,0.18468876,0.13724397,0.00020015532],"about_ca_topic_score_codex":0.012768483,"about_ca_topic_score_gemma":0.020495364,"teacher_disagreement_score":0.06379175,"about_ca_system_score_codex":0.0029021353,"about_ca_system_score_gemma":0.005217667,"threshold_uncertainty_score":0.21340466},"labels":[],"label_agreement":null},{"id":"W4415250377","doi":"10.2139/ssrn.5614285","title":"Understanding JavaScript to TypeScript Migration: Strategies and Quality Assessment","year":2025,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Ontario Tech University","funders":"","keywords":"TypeScript; JavaScript; Code smell; Maintainability; Popularity; Scope (computer science); Source code","score_opus":0.10449999417301452,"score_gpt":0.35348286977426446,"score_spread":0.24898287560124993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415250377","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05084471,0.0006030742,0.92869,0.0020613435,0.00010572988,0.0002081783,0.00022780597,0.01019235,0.007066754],"genre_scores_gemma":[0.4542199,0.00091962505,0.5278066,0.00063180743,0.00009223043,0.00013874692,0.00077786524,0.009484887,0.0059284163],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9917977,0.002588253,0.00065208913,0.0011787972,0.0033050478,0.00047816947],"domain_scores_gemma":[0.96226865,0.019182272,0.0027510044,0.0073168534,0.007969604,0.0005116175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009500383,0.0010068297,0.0007321349,0.0018105626,0.0010184849,0.009358759,0.0029545915,0.0022077973,0.0046575186],"category_scores_gemma":[0.080590986,0.0014353695,0.00091121177,0.0016113535,0.0018322312,0.014176138,0.0030352324,0.0033494493,0.0014606697],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004679317,0.0003098124,0.020495765,0.0009934918,0.00010466949,0.0006408794,0.014290886,0.023093494,0.028090615,0.17637116,0.013319225,0.72182196],"study_design_scores_gemma":[0.00010305357,0.00023368176,0.008056656,0.0010592794,0.00034292665,0.0010319299,0.0061532874,0.46219718,0.10805798,0.29222423,0.120336294,0.00020350429],"about_ca_topic_score_codex":0.005212514,"about_ca_topic_score_gemma":0.005424692,"teacher_disagreement_score":0.009500383,"about_ca_system_score_codex":0.0019739058,"about_ca_system_score_gemma":0.0029936684,"threshold_uncertainty_score":0.050243437},"labels":[],"label_agreement":null},{"id":"W4415312861","doi":"10.1145/3771929","title":"Continuously Learning Bug Locations","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Software bug; Software regression; Software; Code (set theory); Deep learning; Forgetting; Source code; Mean reciprocal rank","score_opus":0.04947276161402145,"score_gpt":0.3231860867865974,"score_spread":0.27371332517257596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415312861","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5342503,0.003954112,0.4340063,0.0015852085,0.00040619762,0.00019927065,0.0026233434,0.019321602,0.0036537189],"genre_scores_gemma":[0.9280525,0.00038265463,0.0649595,0.00028529725,0.000117498756,0.00009218863,0.0033445172,0.00022605126,0.0025398291],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985266,0.00021729454,0.00009571145,0.0007631685,0.00028289217,0.000114296156],"domain_scores_gemma":[0.9939167,0.0030004766,0.00085401675,0.00065267703,0.0012618155,0.0003142693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001409481,0.0017457611,0.0010077351,0.0025858225,0.00036635794,0.0009595298,0.0019975842,0.0012522157,0.0014134645],"category_scores_gemma":[0.011298339,0.0005467289,0.00092426885,0.0012958106,0.0005084722,0.0021471402,0.0014157715,0.0017289396,0.0011337595],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051250384,0.00067851593,0.08024874,0.0004939668,0.00026349002,0.00044673614,0.00027337318,0.23556119,0.008460787,0.0010532326,0.013945974,0.6580615],"study_design_scores_gemma":[0.000026821413,0.0001727863,0.005820498,0.000036674093,0.00005357281,0.00012788399,0.000062605664,0.98693407,0.002750182,0.0023394332,0.0016570938,0.00001834895],"about_ca_topic_score_codex":0.0066604395,"about_ca_topic_score_gemma":0.008770238,"teacher_disagreement_score":0.0066604395,"about_ca_system_score_codex":0.0007993166,"about_ca_system_score_gemma":0.0014096233,"threshold_uncertainty_score":0.013243318},"labels":[],"label_agreement":null},{"id":"W4415323141","doi":"10.1007/s10515-025-00556-y","title":"Graph neural networks for precise bug localization through structural program analysis","year":2025,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Graph; Classifier (UML); Software; Source code; Program analysis; Software bug; Artificial neural network; Debugging","score_opus":0.008905159232018756,"score_gpt":0.28274511318007195,"score_spread":0.2738399539480532,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415323141","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19890663,0.0022540977,0.78082323,0.00089185336,0.00017695947,0.0002496947,0.0018723266,0.011612841,0.0032124077],"genre_scores_gemma":[0.80278677,0.00060525525,0.1886875,0.00027319888,0.00008478211,0.00020124143,0.004214796,0.00024745404,0.0028989064],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99942267,0.00011912228,0.000035356974,0.00022767334,0.00012281681,0.000072298186],"domain_scores_gemma":[0.9979767,0.0010017564,0.0003358332,0.00016401634,0.00045683267,0.00006484711],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007131019,0.0014253603,0.000826809,0.004442336,0.00050352706,0.0007519764,0.0015279927,0.0012463115,0.0017864413],"category_scores_gemma":[0.0040060235,0.00041032734,0.00081079634,0.002274221,0.00046716255,0.0013092889,0.00069442723,0.0012501469,0.00067058695],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035937142,0.0003756839,0.010951653,0.00028771028,0.00016605554,0.0002216885,0.0001425106,0.4755333,0.006192873,0.003624153,0.007982046,0.49416295],"study_design_scores_gemma":[0.0000041124986,0.00001589674,0.0004998648,0.000010097628,0.00001219286,0.0000103401135,0.000011620668,0.9962734,0.0006325337,0.0022474553,0.00027950073,0.0000031239404],"about_ca_topic_score_codex":0.018762967,"about_ca_topic_score_gemma":0.019433197,"teacher_disagreement_score":0.018762967,"about_ca_system_score_codex":0.0013576775,"about_ca_system_score_gemma":0.0009138702,"threshold_uncertainty_score":0.0373075},"labels":[],"label_agreement":null},{"id":"W4415481257","doi":"10.1109/tse.2025.3624631","title":"Contrasting the Hyperparameter Tuning Impact Across Software Defect Prediction Scenarios","year":2025,"lang":"","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Hyperparameter; Software; Hyperparameter optimization; Software quality assurance; Software bug; Scope (computer science); Quality (philosophy)","score_opus":0.015967407007470576,"score_gpt":0.2777652525143255,"score_spread":0.2617978455068549,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415481257","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94948083,0.002195869,0.042347077,0.00077234977,0.00015871933,0.00022742714,0.0006589225,0.0017874922,0.0023713757],"genre_scores_gemma":[0.978438,0.00029571494,0.019380337,0.0001830628,0.00004300483,0.0001160037,0.0011577134,0.00011334306,0.0002727956],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9899809,0.004839719,0.00097517134,0.0021251643,0.0014935879,0.00058548816],"domain_scores_gemma":[0.9223199,0.061142135,0.0037346093,0.007933238,0.003996712,0.00087334757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013249129,0.002026124,0.00082761806,0.0018767365,0.0005737123,0.0017912093,0.0013341994,0.0019277601,0.00045641727],"category_scores_gemma":[0.06672763,0.0005084647,0.0009259802,0.0013718787,0.00103117,0.0029498672,0.0013793211,0.0024071592,0.00024327531],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021534476,0.0013327636,0.074624345,0.00045149843,0.0007635978,0.0002615709,0.00029085297,0.7597412,0.0134532275,0.0017953288,0.0039902492,0.14114186],"study_design_scores_gemma":[0.0002486484,0.0018507919,0.03589935,0.00010853346,0.00027501304,0.00024598508,0.0003737556,0.93270034,0.023073275,0.0034915837,0.0016254337,0.00010728609],"about_ca_topic_score_codex":0.0036198723,"about_ca_topic_score_gemma":0.0025767598,"teacher_disagreement_score":0.013249129,"about_ca_system_score_codex":0.0009889441,"about_ca_system_score_gemma":0.0009017509,"threshold_uncertainty_score":0.070068955},"labels":[],"label_agreement":null},{"id":"W4415665746","doi":"10.1145/3772008.3772019","title":"The First International Workshop on Neuro-Symbolic Software Engineering (NSE 2025)","year":2025,"lang":"en","type":"article","venue":"ACM SIGSOFT Software Engineering Notes","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Multidisciplinary approach; Social software engineering; Software development; Work (physics); Software; Software requirements; Point (geometry); Software construction","score_opus":0.014921288394421762,"score_gpt":0.2594066150306536,"score_spread":0.24448532663623185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415665746","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023864487,0.03387057,0.6764201,0.0502063,0.042215634,0.00064182753,0.0022892456,0.003167556,0.16732438],"genre_scores_gemma":[0.19254625,0.026478142,0.38418898,0.008032234,0.008542225,0.0010361731,0.008214846,0.0035808743,0.36738023],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99534273,0.0015154708,0.00023753608,0.00060958485,0.0017177546,0.0005769037],"domain_scores_gemma":[0.9946515,0.0016848925,0.0001094391,0.00068440085,0.00188572,0.0009840059],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008531116,0.0013839508,0.0009619728,0.0015166901,0.0010624931,0.005627025,0.0020910494,0.0024623745,0.022819016],"category_scores_gemma":[0.010038958,0.000533052,0.0016134175,0.0012361312,0.0018370495,0.003702131,0.0046609333,0.0043296553,0.0074973525],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039516794,0.00033703473,0.0008148641,0.000554306,0.00011651364,0.00029233316,0.0010303138,0.013309674,0.0070929625,0.13228168,0.36658707,0.47718805],"study_design_scores_gemma":[0.000050152954,0.00016332894,0.0006256926,0.00040328153,0.00003104623,0.00018993925,0.00028584263,0.01696964,0.0052997973,0.051422738,0.924515,0.00004344294],"about_ca_topic_score_codex":0.0048252866,"about_ca_topic_score_gemma":0.006011677,"teacher_disagreement_score":0.022819016,"about_ca_system_score_codex":0.0032227829,"about_ca_system_score_gemma":0.006151324,"threshold_uncertainty_score":0.07633722},"labels":[],"label_agreement":null},{"id":"W4415744470","doi":"10.1109/qrs-c65679.2025.00029","title":"Software Defects Prediction: Source Code Is All You Need","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Source code; Software quality; Process (computing); Reliability (semiconductor); Code (set theory); Software; Quality assurance; Static program analysis; Precision and recall","score_opus":0.01832959450004438,"score_gpt":0.2683859607950172,"score_spread":0.2500563662949728,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415744470","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56762815,0.0031393233,0.36879677,0.005089942,0.00043142712,0.00024405148,0.012443231,0.02859327,0.013633822],"genre_scores_gemma":[0.90631866,0.0007662813,0.072229974,0.0004634232,0.00014405334,0.00008222986,0.011959258,0.00085994636,0.007176185],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990369,0.000097712524,0.00005909745,0.00032244594,0.0004128239,0.00007109394],"domain_scores_gemma":[0.99443394,0.002060939,0.0011042921,0.0006890304,0.0015000309,0.00021181491],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007337354,0.00093453773,0.0005082347,0.0023515855,0.00032837194,0.0010829308,0.0007098412,0.00086189416,0.0023398865],"category_scores_gemma":[0.009886986,0.0002759521,0.0004923698,0.0014359115,0.00036012675,0.0024059068,0.0007424633,0.0010184956,0.0020950693],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003984223,0.00030555908,0.22283927,0.00055504293,0.00015623927,0.00051588495,0.00038613615,0.022224301,0.01596467,0.0027747315,0.0443365,0.68954325],"study_design_scores_gemma":[0.000032855325,0.00045355078,0.15403618,0.00048147296,0.00020333314,0.001599211,0.0006555742,0.7259436,0.040201355,0.021682737,0.05458928,0.00012086205],"about_ca_topic_score_codex":0.0048280004,"about_ca_topic_score_gemma":0.0065480317,"teacher_disagreement_score":0.0048280004,"about_ca_system_score_codex":0.00035999445,"about_ca_system_score_gemma":0.00066729565,"threshold_uncertainty_score":0.009599805},"labels":[],"label_agreement":null},{"id":"W4415746151","doi":"10.1109/icsme64153.2025.00051","title":"Does Editing Improve Answer Quality on Stack Overflow? A Data-Driven Investigation","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Readability; Code (set theory); Quality (philosophy); Key (lock); Code review; Software; Collaborative editing","score_opus":0.05128130173242224,"score_gpt":0.342831994137282,"score_spread":0.29155069240485976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415746151","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9941848,0.00046154778,0.0017904183,0.00041180738,0.000017347595,0.00010347027,0.0004995905,0.00014204618,0.0023890212],"genre_scores_gemma":[0.99707687,0.0001178459,0.0014618945,0.00011236122,0.000031474385,0.00008241536,0.00049656397,0.00010745587,0.0005131745],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9602939,0.01753511,0.0034705848,0.0039739744,0.012956649,0.0017698006],"domain_scores_gemma":[0.18484402,0.71155965,0.06636222,0.014636117,0.020480996,0.0021170839],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.04138067,0.0006015372,0.0009942079,0.0045427787,0.0010274432,0.0049096416,0.0016312047,0.0014754765,0.0033722154],"category_scores_gemma":[0.39161736,0.0004677191,0.0016328115,0.004778557,0.0029381616,0.0065370006,0.003369753,0.0022846372,0.0006487246],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024933962,0.0008959194,0.8740267,0.001524437,0.00080858025,0.00041580608,0.022905795,0.0021045364,0.0043272963,0.0013946109,0.0017816261,0.08732127],"study_design_scores_gemma":[0.000048768597,0.0007658273,0.98221856,0.00019088817,0.0003093378,0.00020451283,0.004115392,0.0045412336,0.0036787991,0.0012938677,0.00256665,0.00006632656],"about_ca_topic_score_codex":0.0028509058,"about_ca_topic_score_gemma":0.0034345456,"teacher_disagreement_score":0.99509037,"about_ca_system_score_codex":0.0021196513,"about_ca_system_score_gemma":0.0016211192,"threshold_uncertainty_score":0.21884453},"labels":[],"label_agreement":null},{"id":"W4415746161","doi":"10.1109/icsme64153.2025.00078","title":"From First Use to Final Commit: Studying the Evolution of Multi-CI Service Adoption","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University","funders":"","keywords":"Service (business); Java; Empirical research; Cloud computing","score_opus":0.07298606082396908,"score_gpt":0.30460087645346656,"score_spread":0.23161481562949748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415746161","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99715245,0.0001551791,0.0011655479,0.00028651688,0.0000055345304,0.000020060223,0.0001186415,0.00004122883,0.0010548125],"genre_scores_gemma":[0.9980082,0.000090835005,0.001292305,0.00003486496,0.000005819361,0.000027770035,0.00020003389,0.00002522041,0.00031496407],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99294466,0.0022643688,0.00056358037,0.0013111386,0.0020651908,0.00085112225],"domain_scores_gemma":[0.8964525,0.049384788,0.027012793,0.006109078,0.016029086,0.005011732],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009960854,0.000364746,0.0003997939,0.004333962,0.0009365928,0.0031555344,0.0013024489,0.0010276473,0.0012671973],"category_scores_gemma":[0.08890326,0.00046964712,0.0003609183,0.0045373915,0.0013139154,0.0047213887,0.0025110275,0.001935321,0.00040700796],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017111862,0.0002314789,0.9153777,0.000103097445,0.000060819722,0.00038542837,0.023291348,0.0009226603,0.0013813189,0.0011097855,0.0007541929,0.056211114],"study_design_scores_gemma":[0.000011114959,0.00025091472,0.9670085,0.00008928358,0.000030368892,0.0003505829,0.020556444,0.0070207017,0.0008595625,0.0008274495,0.0029374966,0.00005764531],"about_ca_topic_score_codex":0.017192816,"about_ca_topic_score_gemma":0.019840833,"teacher_disagreement_score":0.017192816,"about_ca_system_score_codex":0.002293063,"about_ca_system_score_gemma":0.0020243118,"threshold_uncertainty_score":0.052678704},"labels":[],"label_agreement":null},{"id":"W4415746321","doi":"10.1109/icsme64153.2025.00077","title":"Can LLMs Write CI? a Study on Automatic Generation of GitHub Actions Configurations","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Executable; Automation; Similarity (geometry); Quality (philosophy); Software; Code (set theory)","score_opus":0.07000254763821248,"score_gpt":0.3420604817036072,"score_spread":0.2720579340653947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415746321","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7176997,0.0029957427,0.12510361,0.0011575724,0.00051439536,0.00094998674,0.018711943,0.11810161,0.014765404],"genre_scores_gemma":[0.7169682,0.000814124,0.19825244,0.00071190065,0.000051457944,0.00074401626,0.069850005,0.0073340023,0.0052739717],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959431,0.0018229081,0.0002775922,0.0010634228,0.0006989317,0.00019401377],"domain_scores_gemma":[0.98185974,0.01084898,0.0011359067,0.00435292,0.0013728922,0.00042960834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004006775,0.001721536,0.0006013305,0.001744648,0.0007029933,0.0017564428,0.0025941853,0.0014255884,0.004117515],"category_scores_gemma":[0.02751347,0.0008017833,0.0009993108,0.0015658292,0.00097038364,0.0033715486,0.0020336835,0.002239718,0.002834459],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025300302,0.0015817245,0.045356967,0.0048324596,0.0003268747,0.0017576913,0.0049492805,0.095736235,0.023310827,0.009288295,0.107452065,0.70287746],"study_design_scores_gemma":[0.00065183145,0.0018974234,0.030531473,0.00089556933,0.00024314805,0.0020969652,0.004036846,0.7573304,0.07362053,0.010919605,0.117473334,0.00030287867],"about_ca_topic_score_codex":0.008615357,"about_ca_topic_score_gemma":0.013944437,"teacher_disagreement_score":0.008615357,"about_ca_system_score_codex":0.0018361283,"about_ca_system_score_gemma":0.002062604,"threshold_uncertainty_score":0.021190107},"labels":[],"label_agreement":null},{"id":"W4415821354","doi":"10.1109/tse.2025.3627891","title":"Causes and Canonicalization of Unreproducible Builds in Java","year":2025,"lang":"","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Artifact (error); Java; Software; Taxonomy (biology); Focus (optics); Identification (biology); Software development; Legacy system","score_opus":0.014029697954729706,"score_gpt":0.2637907972674899,"score_spread":0.2497610993127602,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415821354","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9569225,0.0015486641,0.017020237,0.00062072615,0.00005610615,0.00014519556,0.015708236,0.0037611942,0.0042172475],"genre_scores_gemma":[0.9254166,0.0010799365,0.019269552,0.00024552274,0.00005050439,0.00024261558,0.0507984,0.0011352141,0.0017616496],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9867597,0.0027223215,0.0018995146,0.0027361163,0.0050926832,0.00078962144],"domain_scores_gemma":[0.8934863,0.047383662,0.022436693,0.026514603,0.008815219,0.0013634969],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004678768,0.00048817418,0.00044648562,0.0066070217,0.0013850775,0.0017079617,0.0011554394,0.0009705625,0.0009034631],"category_scores_gemma":[0.047979247,0.0005316551,0.0008102873,0.007904416,0.0014749378,0.0017982329,0.002419295,0.0012667907,0.0005138615],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003105933,0.00019537158,0.8580597,0.0007333464,0.00015980443,0.001815995,0.0046854517,0.006200124,0.004820467,0.0062597333,0.018227087,0.09853227],"study_design_scores_gemma":[0.000055578257,0.0001526745,0.85151976,0.00048473378,0.00018169961,0.0046070814,0.00405222,0.03133981,0.015487609,0.0075559043,0.08443196,0.00013093825],"about_ca_topic_score_codex":0.01060992,"about_ca_topic_score_gemma":0.01861179,"teacher_disagreement_score":0.9953212,"about_ca_system_score_codex":0.0010630742,"about_ca_system_score_gemma":0.0025864055,"threshold_uncertainty_score":0.024743974},"labels":[],"label_agreement":null},{"id":"W4416127924","doi":"10.1002/spe.70030","title":"Exploring the Impacts of Antipatterns on Object‐Oriented, Service‐Oriented, and Mobile‐Oriented Systems","year":2025,"lang":"en","type":"article","venue":"Software Practice and Experience","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi; Concordia University","funders":"","keywords":"Coding (social sciences); Software quality; Program comprehension; Comprehension; Source code; Software; Quality (philosophy); Software maintenance","score_opus":0.031122302989594865,"score_gpt":0.3096959911286585,"score_spread":0.27857368813906364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416127924","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46262127,0.50649345,0.013497348,0.0041592107,0.00033687116,0.00090811163,0.0020074837,0.00009256662,0.009883756],"genre_scores_gemma":[0.88318706,0.10495205,0.009539747,0.00084318133,0.0000874383,0.0005861575,0.00047037806,0.000030520838,0.00030348595],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9428074,0.027002294,0.0127876,0.0032178035,0.013430154,0.00075487717],"domain_scores_gemma":[0.57190704,0.31819674,0.07149438,0.0072003934,0.030090187,0.0011112576],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.042632088,0.00050673494,0.0012151228,0.011143247,0.0006150098,0.0033891874,0.0008833496,0.00072981033,0.0015135155],"category_scores_gemma":[0.18185897,0.00054404,0.003021139,0.014414505,0.001890196,0.004519496,0.0025386154,0.0010013675,0.00013683343],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005370266,0.00016506246,0.24585153,0.15951096,0.009595732,0.00047676617,0.012923104,0.0013452547,0.0014286579,0.006013943,0.0018465582,0.5603054],"study_design_scores_gemma":[0.00025521807,0.002586259,0.54481566,0.31528538,0.028145326,0.0020730037,0.022567078,0.002441674,0.005534023,0.01098613,0.065141276,0.00016898739],"about_ca_topic_score_codex":0.005354525,"about_ca_topic_score_gemma":0.0125993015,"teacher_disagreement_score":0.042632088,"about_ca_system_score_codex":0.002651629,"about_ca_system_score_gemma":0.0077503026,"threshold_uncertainty_score":0.22546273},"labels":[],"label_agreement":null},{"id":"W4416183116","doi":"10.1109/gaclm67198.2025.11232363","title":"Revolutionizing Code Optimization: A Deep Dive into AI-Powered Tools for Software Development","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université TÉLUQ","funders":"","keywords":"Software; Software development; Code (set theory); Software quality; Static program analysis; Software system; Code review; Reliability (semiconductor); Program optimization","score_opus":0.02778767082240828,"score_gpt":0.29772632745908306,"score_spread":0.26993865663667477,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416183116","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27874833,0.031658188,0.5992454,0.01343073,0.0006719625,0.0005180349,0.0019219662,0.031845022,0.04196032],"genre_scores_gemma":[0.32485667,0.007279605,0.6535053,0.001917422,0.00020867938,0.00022808442,0.0031562624,0.006007073,0.0028408463],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934402,0.002483227,0.0004788924,0.0007036974,0.0025601014,0.0003339193],"domain_scores_gemma":[0.9780385,0.011944201,0.0011795823,0.006023503,0.0023191122,0.0004950805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005792986,0.0009458097,0.00044866107,0.002489474,0.0006124536,0.0027951633,0.0020143115,0.00096204435,0.0016082658],"category_scores_gemma":[0.031315412,0.0005650438,0.0007668461,0.0031685126,0.0024384034,0.0058222967,0.0028308777,0.0032097965,0.001184497],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046225276,0.00025755673,0.016465332,0.0022334587,0.00015640837,0.0002641951,0.0015119027,0.032756086,0.020948045,0.036238305,0.030316109,0.85839045],"study_design_scores_gemma":[0.0002843111,0.001106397,0.018433014,0.0014587766,0.00021218366,0.0016050217,0.0015268876,0.27206042,0.043989684,0.12768604,0.5314304,0.00020684939],"about_ca_topic_score_codex":0.0031235446,"about_ca_topic_score_gemma":0.004595183,"teacher_disagreement_score":0.005792986,"about_ca_system_score_codex":0.0010872248,"about_ca_system_score_gemma":0.002556158,"threshold_uncertainty_score":0.030636609},"labels":[],"label_agreement":null},{"id":"W4416236009","doi":"10.48550/arxiv.2511.09000","title":"An insight into the technical debt-fix trade off in software backporting","year":2025,"lang":"","type":"preprint","venue":"ArXiv.org","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Global Institute for Water Security, University of Saskatchewan","keywords":"Technical debt; Eclipse; Python (programming language); Software; Debt; Software maintenance","score_opus":0.03660287445175507,"score_gpt":0.3118980520582694,"score_spread":0.27529517760651434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416236009","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9530381,0.0016621903,0.02866535,0.0032198098,0.00006508991,0.00009165896,0.00045269856,0.00041715647,0.01238797],"genre_scores_gemma":[0.990252,0.0003137583,0.0075889374,0.00016901182,0.000029863671,0.000026817328,0.00029255002,0.00009718064,0.0012297906],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9905144,0.0028626337,0.00091645034,0.0014761877,0.0032367862,0.0009935622],"domain_scores_gemma":[0.9087436,0.047399547,0.02347704,0.0088082645,0.008692693,0.0028788745],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012662815,0.00050278916,0.00045626328,0.0055452976,0.0016154936,0.004065932,0.0012849909,0.0010693204,0.0027612469],"category_scores_gemma":[0.08648013,0.00075286836,0.00050526945,0.0056801923,0.0020077054,0.00855116,0.0034344268,0.0022008135,0.00055666134],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003480325,0.00027258456,0.8018758,0.00028055807,0.00014406518,0.0020073955,0.009969393,0.009288119,0.0067134467,0.02225184,0.004076786,0.14277203],"study_design_scores_gemma":[0.00004813648,0.00047310424,0.86283934,0.0003958408,0.00011085722,0.0037893252,0.014576097,0.046548817,0.0053948723,0.04625006,0.019394785,0.00017878409],"about_ca_topic_score_codex":0.005389021,"about_ca_topic_score_gemma":0.006587843,"teacher_disagreement_score":0.012662815,"about_ca_system_score_codex":0.0023728306,"about_ca_system_score_gemma":0.0017668869,"threshold_uncertainty_score":0.06696814},"labels":[],"label_agreement":null},{"id":"W4416336295","doi":"10.1007/978-3-032-12089-2_38","title":"Exploring the Performance of ML Model Size for Classification in Relation to Energy Consumption","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Commit; Relation (database); Energy consumption; Construct (python library); Variable (mathematics); Reduction (mathematics); Energy (signal processing); Software; Language model; Fraction (chemistry)","score_opus":0.07788542627065352,"score_gpt":0.280680891953476,"score_spread":0.20279546568282247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416336295","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8636035,0.0074539008,0.10884426,0.004292219,0.00037028294,0.00010620891,0.001994267,0.0048012612,0.008534],"genre_scores_gemma":[0.9675784,0.0005068984,0.027148958,0.0002253812,0.00013970275,0.000045021246,0.0020578399,0.0005328729,0.0017648346],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941743,0.003292766,0.00038529938,0.0009078379,0.0009230845,0.00031668934],"domain_scores_gemma":[0.7747359,0.2107541,0.0023348278,0.006821851,0.0045408797,0.0008124804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013929232,0.0016209177,0.001546019,0.0017174492,0.00069794565,0.003958544,0.0018867744,0.0026020228,0.0024551842],"category_scores_gemma":[0.095433965,0.0005931106,0.00094037026,0.0018276949,0.0009596239,0.005349996,0.0014713304,0.0027363305,0.0014018895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007650115,0.001402251,0.09433757,0.00058453454,0.000963707,0.00028314957,0.0005554283,0.49523562,0.014296424,0.0036718135,0.0094547365,0.3715647],"study_design_scores_gemma":[0.00006728706,0.00056921603,0.007323765,0.0000495042,0.00014862843,0.00012337828,0.00026952662,0.98002076,0.007586045,0.003136495,0.0006613414,0.0000440979],"about_ca_topic_score_codex":0.006381586,"about_ca_topic_score_gemma":0.004740588,"teacher_disagreement_score":0.013929232,"about_ca_system_score_codex":0.0010661422,"about_ca_system_score_gemma":0.0011229594,"threshold_uncertainty_score":0.07366568},"labels":[],"label_agreement":null},{"id":"W4416363878","doi":"10.3390/make7040149","title":"Explainable Recommendation of Software Vulnerability Repair Based on Metadata Retrieval and Multifaceted LLMs","year":2025,"lang":"en","type":"article","venue":"Machine Learning and Knowledge Extraction","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Metadata; Context (archaeology); Code (set theory); Vulnerability (computing); Knowledge base; Transparency (behavior); Artifact (error); Robustness (evolution)","score_opus":0.016644298966465746,"score_gpt":0.31910019328848227,"score_spread":0.3024558943220165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416363878","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41160142,0.009780501,0.5341796,0.0026200535,0.00027212448,0.0014112483,0.010383626,0.01808071,0.011670631],"genre_scores_gemma":[0.62049836,0.0011356279,0.3625057,0.0004976943,0.00011112541,0.0004361845,0.011032006,0.00030283973,0.0034804957],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969055,0.0008284186,0.00033223425,0.00074183295,0.0010109857,0.0001809801],"domain_scores_gemma":[0.9891894,0.0059969155,0.0008936132,0.0016204177,0.0019462517,0.0003534024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023073272,0.0014034295,0.0010761239,0.0056548417,0.00066073926,0.0019731855,0.0013981182,0.0016922954,0.0022265732],"category_scores_gemma":[0.022180663,0.0003977864,0.0014193314,0.0027742116,0.00039311082,0.0033388645,0.0012502744,0.0015731466,0.0014126134],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011493722,0.0010845863,0.08016481,0.003009677,0.00079332094,0.0006449288,0.0022550859,0.03142967,0.023807315,0.004521388,0.028964486,0.8221753],"study_design_scores_gemma":[0.00031678323,0.0014976994,0.06559817,0.00094651245,0.0013609335,0.0010973356,0.002883364,0.8346069,0.029598333,0.01675474,0.04494271,0.00039653006],"about_ca_topic_score_codex":0.011485492,"about_ca_topic_score_gemma":0.03551168,"teacher_disagreement_score":0.011485492,"about_ca_system_score_codex":0.0009270068,"about_ca_system_score_gemma":0.0020468545,"threshold_uncertainty_score":0.022837281},"labels":[],"label_agreement":null},{"id":"W4416379844","doi":"10.1007/978-981-96-9203-3_4","title":"Google and Flower Federated Learning Frameworks Comparison in Bug Prediction in Terms of Flexibility and Technical Factors","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in electrical engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Python (programming language); Flexibility (engineering); Software; Federated learning; Reliability (semiconductor); Artificial neural network; Context (archaeology); Data modeling","score_opus":0.00917050062563008,"score_gpt":0.2493320065231766,"score_spread":0.24016150589754653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416379844","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94624054,0.012532577,0.020065412,0.00085183035,0.000435923,0.000113711285,0.0030337544,0.0075838794,0.009142298],"genre_scores_gemma":[0.9665355,0.0013909715,0.02409893,0.0001523058,0.00006955408,0.000043942244,0.0051507885,0.00023679,0.0023211595],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967655,0.00071781815,0.0002717138,0.0005291681,0.0013151978,0.00040059004],"domain_scores_gemma":[0.98858225,0.007188159,0.0005489067,0.0013782522,0.0018025012,0.0004998735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005858628,0.00090846414,0.0009797693,0.004520097,0.00054348493,0.0019123182,0.0012135954,0.0011222017,0.0014338824],"category_scores_gemma":[0.017964644,0.00023076974,0.0013370118,0.0029227785,0.000351276,0.0034626336,0.0010836011,0.00085290225,0.0004925557],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006114922,0.0012871568,0.12051052,0.001136999,0.00084773975,0.0003146925,0.00029782098,0.0684295,0.006930862,0.0057745287,0.027464587,0.76089066],"study_design_scores_gemma":[0.00044791945,0.003405537,0.14626475,0.0005912909,0.0020734037,0.001507441,0.0011726288,0.7861364,0.019777877,0.010960881,0.02740278,0.00025918175],"about_ca_topic_score_codex":0.008385915,"about_ca_topic_score_gemma":0.012131552,"teacher_disagreement_score":0.008385915,"about_ca_system_score_codex":0.0008821797,"about_ca_system_score_gemma":0.0018984888,"threshold_uncertainty_score":0.030983746},"labels":[],"label_agreement":null},{"id":"W4416677538","doi":"10.1109/models67397.2025.00014","title":"MCeT: Behavioral Model Correctness Evaluation using Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Huawei Technologies (Canada)","funders":"","keywords":"Correctness; Sequence diagram; Documentation; Natural language; Automation; Sequence (biology); Hallucinating; Behavioral modeling; Data modeling","score_opus":0.09302292773906319,"score_gpt":0.4021141141634302,"score_spread":0.309091186424367,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416677538","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25746378,0.0002347228,0.6909208,0.0002860825,0.000094065355,0.0004663828,0.0017991886,0.042672448,0.0060625095],"genre_scores_gemma":[0.60093576,0.00007022916,0.39113584,0.00008699827,0.000016766246,0.0003577838,0.003679315,0.0021827389,0.0015345422],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9860169,0.005893466,0.0011013993,0.0011675168,0.005460307,0.0003603605],"domain_scores_gemma":[0.9281676,0.04818778,0.004185011,0.009795477,0.009043608,0.0006205775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0100369165,0.0013726281,0.0006280076,0.0034223315,0.00067265757,0.0017447178,0.001905848,0.0011240093,0.0045690727],"category_scores_gemma":[0.07703894,0.0005577296,0.0010344333,0.0013046097,0.0007973907,0.0036401593,0.002350674,0.0009749528,0.00077416503],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020797818,0.0014328777,0.054172795,0.0013339953,0.0006949128,0.0013072562,0.0022461559,0.4046874,0.061530676,0.034777973,0.022705255,0.4130309],"study_design_scores_gemma":[0.000114474984,0.0002518662,0.0027355691,0.000052918687,0.00004193165,0.00016458037,0.00014375219,0.96127075,0.026157027,0.004923905,0.0041000834,0.000043159514],"about_ca_topic_score_codex":0.0061842594,"about_ca_topic_score_gemma":0.0068230242,"teacher_disagreement_score":0.0100369165,"about_ca_system_score_codex":0.0013432475,"about_ca_system_score_gemma":0.0020190955,"threshold_uncertainty_score":0.053080916},"labels":[],"label_agreement":null},{"id":"W4416828803","doi":"10.1016/j.jss.2025.112714","title":"CREST: Improving interpretability and effectiveness of troubleshooting at ericsson through criterion-Specific trouble report retrieval","year":2025,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Ericsson (Canada)","funders":"Telefonaktiebolaget LM Ericsson; Carleton University","keywords":"Troubleshooting; Interpretability; Relevance (law); Software; Volume (thermodynamics); Data retrieval; Human–computer information retrieval","score_opus":0.014749764668559618,"score_gpt":0.2766366461209478,"score_spread":0.2618868814523882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416828803","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2191583,0.0014640676,0.7326926,0.0006765943,0.00016275998,0.0005356647,0.0015019971,0.03904708,0.004760976],"genre_scores_gemma":[0.6487065,0.0005350302,0.33737186,0.00034039872,0.00008826243,0.00020670121,0.006287988,0.0010645249,0.0053987815],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.996807,0.000743243,0.00026436348,0.00069079717,0.0012291495,0.00026540074],"domain_scores_gemma":[0.9932175,0.0025676398,0.00062741846,0.0012447269,0.0020928974,0.00024980438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030973963,0.0022930074,0.0015077228,0.0035539425,0.00050571223,0.002665495,0.0026142248,0.0017152702,0.0021852925],"category_scores_gemma":[0.016179265,0.0004058753,0.0013625654,0.001448841,0.00072601065,0.0029324049,0.002334557,0.0015869046,0.0025723665],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001519746,0.00059992634,0.013646883,0.000678344,0.00028691423,0.0010802131,0.0012778543,0.21708676,0.04075024,0.0038438116,0.020895055,0.6983343],"study_design_scores_gemma":[0.000076894066,0.00036944583,0.0028280076,0.000041035833,0.00007730487,0.00027177882,0.00034848886,0.96717405,0.0210869,0.0028228287,0.0048204605,0.000082758954],"about_ca_topic_score_codex":0.018324092,"about_ca_topic_score_gemma":0.014753538,"teacher_disagreement_score":0.018324092,"about_ca_system_score_codex":0.0011513063,"about_ca_system_score_gemma":0.0026333893,"threshold_uncertainty_score":0.03643483},"labels":[],"label_agreement":null},{"id":"W4416962840","doi":"10.1109/pst65910.2025.11268830","title":"FragmentFool: Fragment-based Adversarial Perturbation for Graph Neural Network-based Vulnerability Detection","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Source code; Robustness (evolution); Graph; Software; Adversarial system; Vulnerability assessment","score_opus":0.017044019332751986,"score_gpt":0.27549739519579425,"score_spread":0.2584533758630423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416962840","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08033951,0.00046736316,0.90879947,0.00041999394,0.00014569181,0.00016972293,0.00020274911,0.0068114204,0.0026441303],"genre_scores_gemma":[0.88089705,0.00020892796,0.11570931,0.00031458493,0.000033863023,0.00011193586,0.0003405785,0.0002412257,0.0021426678],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993393,0.00017721,0.000020862994,0.00014537164,0.00024087114,0.00007636331],"domain_scores_gemma":[0.998073,0.0010623088,0.00021069277,0.00036483794,0.00022286724,0.000066287685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00091220596,0.0009668326,0.0005611053,0.000775565,0.00035650915,0.00041735318,0.0012535099,0.00089092675,0.001533958],"category_scores_gemma":[0.005225564,0.00024084603,0.0004642265,0.0004019011,0.0010518263,0.0014695934,0.0014743045,0.0015034425,0.0003207172],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000330789,0.00011688151,0.0024052286,0.000109047614,0.00011234426,0.00027138364,0.00009228989,0.8063776,0.020585453,0.009665218,0.005154967,0.15477881],"study_design_scores_gemma":[0.000004852109,0.000048532132,0.00018239819,0.0000042674424,0.0000073230785,0.000051912288,0.0000063890816,0.9921029,0.004087104,0.0031173776,0.00038045098,0.000006528567],"about_ca_topic_score_codex":0.002409807,"about_ca_topic_score_gemma":0.002850941,"teacher_disagreement_score":0.002409807,"about_ca_system_score_codex":0.0008015492,"about_ca_system_score_gemma":0.00056536944,"threshold_uncertainty_score":0.005815685},"labels":[],"label_agreement":null},{"id":"W4417003938","doi":"10.1109/uemcon67449.2025.11267548","title":"Evaluating Source Code Embeddings From LLMS for Educational Downstream Tasks","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Source code; Automatic summarization; Code review; Natural language; Code (set theory); Embedding; Downstream (manufacturing); Paragraph","score_opus":0.04525139358938427,"score_gpt":0.39129531390931643,"score_spread":0.34604392031993214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417003938","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7775912,0.0018507423,0.19157653,0.0010360586,0.00036309593,0.00032847514,0.0036653192,0.019062085,0.004526538],"genre_scores_gemma":[0.8640132,0.00038992392,0.116297774,0.00023892458,0.00005987252,0.00024137512,0.015204694,0.0005872717,0.002966972],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986035,0.00052482815,0.0001299873,0.0004069065,0.00023729802,0.00009754297],"domain_scores_gemma":[0.9935134,0.0040491484,0.00034468822,0.00084292685,0.0009907754,0.0002591481],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018843542,0.0019000815,0.00044156282,0.0013749413,0.00033018328,0.0010127125,0.0010675885,0.00155379,0.0018623449],"category_scores_gemma":[0.017131705,0.00036033746,0.00077442697,0.0009126029,0.0004944086,0.0023496053,0.0014879693,0.0018518778,0.0015246419],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013971445,0.0010740553,0.024031574,0.0008684542,0.00031536043,0.0003031138,0.00041892938,0.36753693,0.021319922,0.002372186,0.015204876,0.5651574],"study_design_scores_gemma":[0.000053032683,0.00032120597,0.0021761875,0.000047354162,0.00004645947,0.000069992595,0.00015785714,0.975802,0.016811118,0.0027170295,0.0017791427,0.00001862727],"about_ca_topic_score_codex":0.0038067035,"about_ca_topic_score_gemma":0.00590794,"teacher_disagreement_score":0.0038067035,"about_ca_system_score_codex":0.000888531,"about_ca_system_score_gemma":0.00092766114,"threshold_uncertainty_score":0.009965539},"labels":[],"label_agreement":null},{"id":"W4417339614","doi":"10.1145/3785001","title":"An Empirical Study of Self-Admitted Technical Debt in Machine Learning Software","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University","funders":"","keywords":"Technical debt; Context (archaeology); Empirical research; Code smell; Pipeline (software); Source code; Code (set theory)","score_opus":0.05521184893937985,"score_gpt":0.3650327064973448,"score_spread":0.30982085755796496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417339614","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9979715,0.00012990068,0.0007750765,0.00017932366,0.00000386875,0.000021605234,0.00020548268,0.000032324202,0.0006808942],"genre_scores_gemma":[0.9983571,0.00009638168,0.00073605415,0.00004500346,0.000009989183,0.00002461752,0.00037897422,0.000023264916,0.00032856775],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9905129,0.002843995,0.0012801408,0.0014412809,0.0032547133,0.0006670812],"domain_scores_gemma":[0.6949841,0.14977771,0.103417456,0.013817937,0.031289227,0.006713547],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009979682,0.00029113403,0.00025353802,0.0043814005,0.0009941591,0.0016892145,0.0008779667,0.0007960689,0.00129067],"category_scores_gemma":[0.14200312,0.00042643427,0.00025092278,0.004446433,0.0019446217,0.004505402,0.0023748598,0.0017356375,0.00045924287],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000978913,0.00010403912,0.97339284,0.00010649709,0.000030166608,0.00028605125,0.006301539,0.00041266106,0.0006705216,0.00049023714,0.0007070336,0.017400509],"study_design_scores_gemma":[0.000007431724,0.00022722855,0.9825613,0.00011244317,0.000015395675,0.0006947956,0.0074303816,0.005000664,0.0007635559,0.0006499116,0.0025036486,0.00003318113],"about_ca_topic_score_codex":0.0029382005,"about_ca_topic_score_gemma":0.0036806588,"teacher_disagreement_score":0.99002033,"about_ca_system_score_codex":0.0011733745,"about_ca_system_score_gemma":0.0009962147,"threshold_uncertainty_score":0.052778244},"labels":[],"label_agreement":null},{"id":"W4417438984","doi":"10.1109/ms.2025.3644251","title":"Understanding Prompt Management in GitHub Repositories: A Call for Best Practices","year":2025,"lang":"","type":"article","venue":"IEEE Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Kingston Health Sciences Centre; Arup Group (Canada); Queen's University","funders":"","keywords":"Readability; Best practice; Maintainability; Usability; Quality (philosophy); Software quality; Natural language","score_opus":0.13182059768940274,"score_gpt":0.36156596925184753,"score_spread":0.2297453715624448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417438984","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23527233,0.019274272,0.5040889,0.2166762,0.0006351441,0.001589017,0.0020027398,0.009300945,0.01116043],"genre_scores_gemma":[0.34225184,0.006884515,0.6392166,0.005820224,0.0002613102,0.00089257944,0.002045515,0.0013768865,0.0012506129],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9197662,0.04445302,0.00884043,0.01145103,0.012150497,0.0033388927],"domain_scores_gemma":[0.6108855,0.24675031,0.030346947,0.06390313,0.041337494,0.006776462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11996903,0.0016506802,0.0023408763,0.017566191,0.005765377,0.03367293,0.00925341,0.006213582,0.0022271834],"category_scores_gemma":[0.30880457,0.0025361786,0.001777475,0.013809241,0.013240678,0.07711914,0.013206837,0.01093163,0.0012176475],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022054104,0.00087184255,0.1101641,0.0053859726,0.00024739536,0.0010679796,0.18595253,0.004554319,0.0073490203,0.043254748,0.021571146,0.6193603],"study_design_scores_gemma":[0.0001234494,0.00047306862,0.09420099,0.016058046,0.00041647727,0.0021909273,0.35660926,0.03578193,0.011367102,0.2673516,0.214564,0.00086321164],"about_ca_topic_score_codex":0.027764862,"about_ca_topic_score_gemma":0.041658983,"teacher_disagreement_score":0.11996903,"about_ca_system_score_codex":0.010626054,"about_ca_system_score_gemma":0.020960318,"threshold_uncertainty_score":0.6344645},"labels":[],"label_agreement":null},{"id":"W4417443056","doi":"10.1016/j.jss.2025.112746","title":"A search-based file recommendation approach for infrastructure-as-code evolution","year":2025,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Hôpital Notre-Dame; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software; Server; Path (computing); Genetic algorithm; Similarity (geometry); Space (punctuation); Software evolution","score_opus":0.015030484004162811,"score_gpt":0.27116887446438687,"score_spread":0.2561383904602241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417443056","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1356342,0.0024341058,0.830232,0.0010314941,0.00036504955,0.0008204362,0.0030717528,0.020361328,0.0060496414],"genre_scores_gemma":[0.49106884,0.0005175996,0.494131,0.0003962621,0.0001922397,0.00021876316,0.0042352653,0.00048022473,0.008759821],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99831724,0.00021839878,0.0001146641,0.00043881484,0.00073079224,0.00018011226],"domain_scores_gemma":[0.9965167,0.0012543353,0.0002133449,0.0005871376,0.0012063952,0.0002221313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00093135505,0.0009986279,0.0018927384,0.005551512,0.0011370612,0.0013314246,0.002925854,0.0021657785,0.003616051],"category_scores_gemma":[0.0055188863,0.000403571,0.0014694892,0.0048500863,0.00037166895,0.0023800295,0.0012753991,0.0010150403,0.0015952984],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071069907,0.0010538072,0.014174408,0.00041464495,0.00035841984,0.00046490986,0.00023956764,0.05746186,0.024565004,0.0030880442,0.027938267,0.8695304],"study_design_scores_gemma":[0.000039609382,0.00013677187,0.0022367802,0.000013902981,0.00010074101,0.00022726186,0.00010393243,0.98827577,0.004388865,0.0020128086,0.0024333892,0.00003017579],"about_ca_topic_score_codex":0.02518257,"about_ca_topic_score_gemma":0.055546515,"teacher_disagreement_score":0.02518257,"about_ca_system_score_codex":0.00081796636,"about_ca_system_score_gemma":0.0018449937,"threshold_uncertainty_score":0.050072014},"labels":[],"label_agreement":null},{"id":"W48382302","doi":"","title":"Design steps for developing software measurement standard etalons for ISO 19761 (COSMIC-FFP)","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software measurement; Software; Computer science; Comparability; Measure (data warehouse); Systems engineering; Reliability engineering; Software development; Software quality; Engineering; Operating system; Mathematics; Database","score_opus":0.11121271045268369,"score_gpt":0.3186180139324348,"score_spread":0.20740530347975109,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W48382302","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009666165,0.00030907398,0.96100247,0.0017500417,0.00036720076,0.011431514,0.00042896846,0.0039704703,0.011074162],"genre_scores_gemma":[0.01620025,0.00015702337,0.9758716,0.00029950874,0.00002803635,0.004295436,0.00088098797,0.00040216307,0.0018649512],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9572161,0.015641114,0.0070732087,0.0023738942,0.016666526,0.0010292325],"domain_scores_gemma":[0.8741509,0.026582379,0.009392021,0.0114926975,0.076658525,0.0017234336],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0672315,0.0024094253,0.0011762136,0.008537579,0.0025505559,0.006852242,0.0043477146,0.003070018,0.0045652683],"category_scores_gemma":[0.10452293,0.0019311678,0.0017876091,0.003090416,0.002339094,0.0057598585,0.0029087977,0.004811677,0.00426054],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005707293,0.002019906,0.016160171,0.0034360774,0.00021785982,0.0014269027,0.009474727,0.02056792,0.05572773,0.1507007,0.038780186,0.70091707],"study_design_scores_gemma":[0.0005628072,0.0026603441,0.021748314,0.0072074677,0.0005251857,0.0014790738,0.008945697,0.07702968,0.13028094,0.08393221,0.66481984,0.00080845226],"about_ca_topic_score_codex":0.0073431837,"about_ca_topic_score_gemma":0.007018199,"teacher_disagreement_score":0.0672315,"about_ca_system_score_codex":0.0054822112,"about_ca_system_score_gemma":0.024085654,"threshold_uncertainty_score":0.3555584},"labels":[],"label_agreement":null},{"id":"W48599167","doi":"","title":"Explaining Product Release Planning Results Using Concept Analysis.","year":2008,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Black box; Sensitivity (control systems); Product (mathematics); Computer science; New product development; Series (stratigraphy); Data mining; Engineering; Mathematics; Artificial intelligence","score_opus":0.02839802575744354,"score_gpt":0.2640647064541355,"score_spread":0.23566668069669194,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W48599167","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.096985124,0.0007066016,0.8892925,0.0008653122,0.00005774156,0.0006787626,0.0022328736,0.0018826921,0.007298382],"genre_scores_gemma":[0.32474092,0.00046102263,0.6697246,0.000114756396,0.000021184598,0.0006037047,0.003242329,0.00017776649,0.0009137267],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99785584,0.0010154785,0.00013108786,0.0002702705,0.0006663219,0.000060877104],"domain_scores_gemma":[0.9764845,0.020645034,0.0010337388,0.000756628,0.0010016331,0.000078481426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049028345,0.0015577889,0.0002979539,0.002880955,0.00039982286,0.0017681674,0.001161966,0.00075831264,0.0075791273],"category_scores_gemma":[0.021955067,0.00032236672,0.0010024265,0.0016426054,0.00081655406,0.0027760603,0.0009742452,0.00087947544,0.0008045086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010512971,0.00066283956,0.022599801,0.00551374,0.0005344398,0.0023128556,0.009392569,0.28437868,0.021317285,0.117871776,0.010486496,0.5238783],"study_design_scores_gemma":[0.000339305,0.0010436215,0.012286318,0.0018241302,0.00029922495,0.0013572678,0.0034136523,0.78211975,0.032986503,0.13105707,0.033077605,0.00019562941],"about_ca_topic_score_codex":0.0016243082,"about_ca_topic_score_gemma":0.0012287193,"teacher_disagreement_score":0.0075791273,"about_ca_system_score_codex":0.0009789617,"about_ca_system_score_gemma":0.0013595232,"threshold_uncertainty_score":0.025929034},"labels":[],"label_agreement":null},{"id":"W585036856","doi":"","title":"Quality engineering process for the program design phase of a generic software life cycle","year":2005,"lang":"en","type":"book-chapter","venue":"British Computer Society eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Software development process; Software engineering; Software quality; Software construction; Personal software process; Software quality analyst; Engineering; Systems engineering; Process (computing); Software quality control; Software development; Systems development life cycle; Quality (philosophy); Computer science; Software","score_opus":0.058255713720377536,"score_gpt":0.3105826120823168,"score_spread":0.2523268983619393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W585036856","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036692815,0.00031987534,0.9834958,0.00034989647,0.00004103391,0.00026982094,0.000032747765,0.00043045328,0.011391093],"genre_scores_gemma":[0.09689119,0.0011448595,0.8892879,0.000307058,0.00004667253,0.000511209,0.00023965746,0.00027386608,0.011297618],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99685466,0.0009039297,0.00017086149,0.00031672657,0.0015763242,0.00017740949],"domain_scores_gemma":[0.99726737,0.00097412226,0.00029114902,0.0005922609,0.0007753045,0.00009975335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032346514,0.00061985786,0.00033328292,0.0010869168,0.0009554893,0.0024226648,0.0012032368,0.0014182857,0.004376966],"category_scores_gemma":[0.0069763213,0.00036322753,0.00091006,0.00085779245,0.0015678619,0.002325261,0.0014943214,0.0019922708,0.0014069497],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006233122,0.00018782646,0.0010039728,0.0007081769,0.000026079395,0.00029464855,0.0024685941,0.029022055,0.009924769,0.61813384,0.006368438,0.33179927],"study_design_scores_gemma":[0.00011987519,0.00052683876,0.0017383442,0.0011350913,0.000098036726,0.0013122805,0.0007173194,0.13959688,0.021585692,0.35267496,0.4803904,0.00010421932],"about_ca_topic_score_codex":0.0018397087,"about_ca_topic_score_gemma":0.0016729429,"teacher_disagreement_score":0.004376966,"about_ca_system_score_codex":0.0012916101,"about_ca_system_score_gemma":0.0030440034,"threshold_uncertainty_score":0.017106712},"labels":[],"label_agreement":null},{"id":"W585295213","doi":"","title":"Identification et localisation des preoccupations fonctionnelles dans du code legataire java","year":2012,"lang":"fr","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Java; Object-oriented programming; Programming language; Aspect-oriented programming; Separation of concerns; Software engineering; Identification (biology); Software","score_opus":0.04136052623276992,"score_gpt":0.31474851727008757,"score_spread":0.27338799103731765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W585295213","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4238193,0.00112804,0.5585903,0.0005191404,0.00017037931,0.00020985573,0.00033126934,0.0060917474,0.009139886],"genre_scores_gemma":[0.66098785,0.0009678884,0.31903863,0.0001907491,0.000097906355,0.00025057117,0.0011909795,0.0017663065,0.015509015],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99496967,0.001103032,0.00033132045,0.0013166643,0.0018996245,0.00037973144],"domain_scores_gemma":[0.9890215,0.0064931815,0.0011063697,0.0012399327,0.0018276133,0.00031144972],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023738297,0.0011224395,0.00077879213,0.003474552,0.0013429731,0.0043583084,0.00070336217,0.0017568684,0.002244529],"category_scores_gemma":[0.012680647,0.00079648016,0.0011360018,0.0016423818,0.00224786,0.0026332815,0.0010186557,0.0021951674,0.0009522049],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012215513,0.0004619313,0.10031045,0.0008897041,0.00023286605,0.0043987865,0.018718401,0.0248967,0.17292772,0.06494343,0.0051052566,0.60589325],"study_design_scores_gemma":[0.00017908428,0.00084133627,0.2243164,0.0005486937,0.00030216196,0.009931334,0.006007638,0.39448836,0.19786607,0.017346937,0.14775757,0.0004144327],"about_ca_topic_score_codex":0.016669136,"about_ca_topic_score_gemma":0.013393958,"teacher_disagreement_score":0.016669136,"about_ca_system_score_codex":0.0017965228,"about_ca_system_score_gemma":0.0017797118,"threshold_uncertainty_score":0.033144176},"labels":[],"label_agreement":null},{"id":"W61949798","doi":"","title":"SHriMP Views: An Interactive Environment for Exploring Multiple Hierarchical Views of a Java Program","year":2001,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Java; Documentation; Visualization; Software engineering; Interoperability; Software; Human–computer interaction; World Wide Web; Programming language; Artificial intelligence","score_opus":0.13610935573491853,"score_gpt":0.346068443486984,"score_spread":0.20995908775206548,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W61949798","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011892052,0.00022761375,0.9066077,0.00022773053,0.00005413126,0.00013732658,0.0009970942,0.06934683,0.010509523],"genre_scores_gemma":[0.12512748,0.0005530847,0.8512494,0.00028555552,0.000048697104,0.00049791846,0.0024765835,0.009507871,0.010253441],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99967253,0.00009823903,0.00001851502,0.00005475903,0.00011548471,0.000040570005],"domain_scores_gemma":[0.99870014,0.0008576966,0.00005190965,0.00016376359,0.000078702185,0.00014776208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00074188435,0.0008279739,0.00043970734,0.00081989495,0.00042196488,0.0012884695,0.0015368396,0.0009641456,0.021296777],"category_scores_gemma":[0.002340531,0.00068898743,0.00072653807,0.00043154633,0.00053674483,0.003121836,0.00322108,0.0018042739,0.0032381127],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00170341,0.00042828926,0.004419883,0.0015215817,0.00013397078,0.0020564138,0.0049115266,0.015060639,0.1778641,0.046259742,0.12523091,0.6204097],"study_design_scores_gemma":[0.00046804632,0.0006565862,0.0054894555,0.0004938287,0.00012733386,0.0029009366,0.0008770523,0.24290806,0.09951612,0.041797396,0.60441005,0.00035512913],"about_ca_topic_score_codex":0.0012726457,"about_ca_topic_score_gemma":0.0029198856,"teacher_disagreement_score":0.021296777,"about_ca_system_score_codex":0.00022690657,"about_ca_system_score_gemma":0.0005320056,"threshold_uncertainty_score":0.071244776},"labels":[],"label_agreement":null},{"id":"W63450075","doi":"","title":"Testing Models or Fitting Models? Identifying Model Misspecification in PLS","year":2010,"lang":"en","type":"article","venue":"International Conference on Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Partial least squares regression; Structural equation modeling; Statistical model; Latent variable; Computer science; Data modeling; Set (abstract data type); Econometrics; Statistical hypothesis testing; Data set; Data mining; Statistics; Machine learning; Artificial intelligence; Mathematics","score_opus":0.19584403487286975,"score_gpt":0.34112173941897794,"score_spread":0.1452777045461082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W63450075","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15428065,0.0017656139,0.82351834,0.014059435,0.00027250077,0.00023829516,0.000363028,0.0010229412,0.004479202],"genre_scores_gemma":[0.8168028,0.0007188124,0.17945659,0.001255743,0.0001946952,0.00031060542,0.0004023439,0.0002928474,0.00056556176],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.8753105,0.10190744,0.0041855066,0.007150628,0.009913907,0.00153205],"domain_scores_gemma":[0.4597771,0.47534236,0.0231033,0.029530738,0.010275822,0.0019707454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.09820573,0.00304007,0.004270594,0.0034956464,0.001645592,0.0060360846,0.0046226634,0.0056952396,0.0042497274],"category_scores_gemma":[0.46122068,0.0015150547,0.0027825057,0.0050754747,0.008046585,0.018105231,0.0040608686,0.0076751336,0.0008125752],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012697553,0.0007080687,0.116121784,0.0022254835,0.0037007462,0.0018271534,0.010395502,0.29201964,0.0032036267,0.3152841,0.008002019,0.2452422],"study_design_scores_gemma":[0.00020629761,0.000703762,0.011432604,0.00067150017,0.00042312447,0.00064159185,0.0027949358,0.47310337,0.0031273088,0.5019673,0.0046809777,0.00024714973],"about_ca_topic_score_codex":0.004023848,"about_ca_topic_score_gemma":0.003533985,"teacher_disagreement_score":0.09820573,"about_ca_system_score_codex":0.0022221971,"about_ca_system_score_gemma":0.003910091,"threshold_uncertainty_score":0.51936775},"labels":[],"label_agreement":null},{"id":"W636190832","doi":"","title":"Evaluating design decay during software evolution","year":2012,"lang":"en","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Software design; Computer science; Software; Software engineering; Software maintenance; Software development; Software system; Programming language","score_opus":0.056879198858096426,"score_gpt":0.3467680889504366,"score_spread":0.28988889009234015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W636190832","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9575019,0.00078232266,0.037218723,0.0001339525,0.000044372217,0.00029675785,0.00036292378,0.00059912266,0.0030599318],"genre_scores_gemma":[0.94847286,0.000345581,0.048274603,0.000081485945,0.000018162613,0.0002590117,0.0011908511,0.00013845808,0.0012190122],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.98382807,0.003896408,0.0015421641,0.0017778772,0.008303681,0.00065189047],"domain_scores_gemma":[0.87138695,0.06722262,0.024713038,0.007997976,0.025708193,0.0029711183],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012416562,0.0009800046,0.0007373499,0.005692054,0.00070384966,0.0020924718,0.0011486663,0.0012616544,0.0009535072],"category_scores_gemma":[0.1101404,0.000505432,0.00074413605,0.0026017495,0.0010676378,0.003323639,0.0018088769,0.0013827183,0.00038095604],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001977039,0.0011805616,0.47393864,0.0013476805,0.00041299374,0.0005130133,0.008730284,0.026220117,0.044827025,0.0026788309,0.0015128021,0.43666112],"study_design_scores_gemma":[0.00015402304,0.007846869,0.7298204,0.00050053484,0.00069080695,0.0015794169,0.006607344,0.16857556,0.06348615,0.0067986874,0.013561758,0.00037846997],"about_ca_topic_score_codex":0.0025917396,"about_ca_topic_score_gemma":0.0025384282,"teacher_disagreement_score":0.012416562,"about_ca_system_score_codex":0.0016231292,"about_ca_system_score_gemma":0.0012930281,"threshold_uncertainty_score":0.06566584},"labels":[],"label_agreement":null},{"id":"W641080781","doi":"","title":"Proceedings : 18th IEEE International Conference on Automated Software Engineering, Montreal, Quebec, Canada, October 6 to 10, 2003","year":2003,"lang":"en","type":"book","venue":"IEEE Computer Society eBooks","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Software engineering; Software; Computer science; Software construction; Collaborative software; Social software engineering; Software system; Software development; Operating system","score_opus":0.018784986087293783,"score_gpt":0.2399565165533756,"score_spread":0.22117153046608182,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W641080781","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016295606,0.057287328,0.20861079,0.017613048,0.031365328,0.00094814127,0.009180289,0.012312146,0.64638734],"genre_scores_gemma":[0.017576316,0.014195692,0.027730307,0.0011106578,0.0010763675,0.00017451443,0.007396397,0.0011163179,0.92962337],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993,0.000087210035,0.000023641358,0.00008602681,0.00040019976,0.00010296375],"domain_scores_gemma":[0.9976954,0.0001833746,0.000037524704,0.00019744628,0.0015303668,0.00035583097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016136249,0.0015202456,0.0012388941,0.0013517914,0.0017597142,0.003979079,0.0021239927,0.0010750489,0.18495744],"category_scores_gemma":[0.0023677645,0.00062130165,0.00046147584,0.0017179153,0.001190707,0.0019293939,0.0011069018,0.0017317844,0.062915556],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008429982,0.00008388232,0.0007910958,0.00012145797,0.000027817905,0.00010214316,0.00014146377,0.0005028527,0.0018481286,0.0026202567,0.8229631,0.17071353],"study_design_scores_gemma":[0.000030340989,0.000038031892,0.0025462948,0.00012593428,0.000042139,0.00013050722,0.00018173145,0.0038428016,0.0011165761,0.002826994,0.98908746,0.000031144507],"about_ca_topic_score_codex":0.29204956,"about_ca_topic_score_gemma":0.50864035,"teacher_disagreement_score":0.29204956,"about_ca_system_score_codex":0.00404712,"about_ca_system_score_gemma":0.007507855,"threshold_uncertainty_score":0.61874425},"labels":[],"label_agreement":null},{"id":"W648401726","doi":"","title":"Integration de la visualisation a multiples vues pour le developpement du logiciel","year":2011,"lang":"fr","type":"dissertation","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Visualization; Software visualization; Software; Software development; Context (archaeology); Software engineering; Human–computer interaction; Software construction; Artificial intelligence; Programming language; Geography","score_opus":0.056880612810804164,"score_gpt":0.32449197295901805,"score_spread":0.2676113601482139,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W648401726","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017527454,0.00151302,0.97026676,0.00063998456,0.00012155231,0.00005017018,0.000073272604,0.004196125,0.005611771],"genre_scores_gemma":[0.15428439,0.0020329403,0.83711624,0.000119813725,0.00013465571,0.00013138009,0.00039338754,0.0010748501,0.004712343],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962411,0.001711157,0.00018125522,0.00044331938,0.0012732939,0.00014997138],"domain_scores_gemma":[0.99548256,0.0026302147,0.00019523507,0.0008652568,0.0006277127,0.00019898418],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035351536,0.0017512485,0.0011250109,0.0023535185,0.0006562909,0.005361674,0.00094850385,0.0014455114,0.0058667175],"category_scores_gemma":[0.009902172,0.0008103545,0.0015147561,0.0015121684,0.0013156773,0.0052267015,0.002982768,0.0034684564,0.0013745512],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055202557,0.00042481767,0.002649677,0.0006718429,0.00022472476,0.0005535918,0.0042802687,0.02202972,0.084642425,0.07878973,0.00628015,0.79890114],"study_design_scores_gemma":[0.00046608926,0.0010758747,0.0127314795,0.00093689235,0.00039649673,0.0027570338,0.0015015057,0.49674228,0.08495286,0.09197943,0.3059574,0.0005025742],"about_ca_topic_score_codex":0.0032164964,"about_ca_topic_score_gemma":0.002429611,"teacher_disagreement_score":0.0058667175,"about_ca_system_score_codex":0.0007386443,"about_ca_system_score_gemma":0.0009606192,"threshold_uncertainty_score":0.019626081},"labels":[],"label_agreement":null},{"id":"W649645005","doi":"","title":"Proceedings of the Fourth European Conference on Software Maintenance and Reengineering, Reengineering Week Zurich, University of Zurich, Switzerland, February 29-March 3, 2000","year":2000,"lang":"en","type":"article","venue":"Medical Entomology and Zoology","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Business process reengineering; Work (physics); Subject (documents); Software maintenance; Engineering management; Engineering; Architecture; Library science; Software; Computer science; Software development; Operations management; History; Archaeology","score_opus":0.009679920454952043,"score_gpt":0.2023249150217682,"score_spread":0.19264499456681616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W649645005","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043823794,0.25128216,0.24612635,0.032519583,0.057382252,0.0012049886,0.0054092673,0.0098388735,0.3524127],"genre_scores_gemma":[0.05157932,0.08591112,0.05102432,0.0023465687,0.006479652,0.00042596087,0.012062753,0.00242153,0.78774875],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986852,0.00026714167,0.00009510371,0.0002740807,0.0005644408,0.00011409616],"domain_scores_gemma":[0.9973049,0.0009225571,0.00013381516,0.00037491543,0.00060482003,0.0006589145],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002773134,0.0011245073,0.001251886,0.0023019055,0.0008678226,0.0043088123,0.001410339,0.001321206,0.082743764],"category_scores_gemma":[0.004517236,0.00056241447,0.00054005225,0.0022478965,0.0010635696,0.003476859,0.0016705961,0.0021627808,0.026612334],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024754717,0.00015028348,0.00087631535,0.0004556269,0.000058212212,0.00015321276,0.0005385709,0.00075308105,0.0032597233,0.0058470406,0.5489428,0.4387176],"study_design_scores_gemma":[0.000042840773,0.00009083735,0.0020379636,0.00031390376,0.00003169735,0.00026144084,0.000156665,0.0011888107,0.00095234823,0.0026472376,0.9922505,0.000025822115],"about_ca_topic_score_codex":0.0040982496,"about_ca_topic_score_gemma":0.0072654015,"teacher_disagreement_score":0.082743764,"about_ca_system_score_codex":0.0011901481,"about_ca_system_score_gemma":0.0018674034,"threshold_uncertainty_score":0.27680546},"labels":[],"label_agreement":null},{"id":"W649946739","doi":"10.71781/10275","title":"Système symbolique de création de résumés de mise à jour","year":2009,"lang":"fr","type":"dissertation","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Art; Humanities","score_opus":0.047059235686457396,"score_gpt":0.35930702440488,"score_spread":0.3122477887184226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W649946739","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027841188,0.0011688307,0.8541988,0.0007331794,0.0007915442,0.0004567511,0.0026405435,0.08880783,0.023361387],"genre_scores_gemma":[0.29805475,0.0013541264,0.63452387,0.00041959842,0.0004492242,0.000926127,0.004152466,0.0046452996,0.055474564],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99826556,0.00036535022,0.000145813,0.00055167923,0.0005669682,0.00010472312],"domain_scores_gemma":[0.9947725,0.0025099404,0.00029005692,0.000932415,0.0011304441,0.00036465214],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026612545,0.0012135381,0.0015466397,0.002730319,0.0016164206,0.0069122417,0.0021129327,0.0016415766,0.022036009],"category_scores_gemma":[0.014096791,0.00077778613,0.0010036243,0.001215473,0.0017498628,0.004673107,0.002557307,0.001964869,0.008158583],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00332163,0.0003983047,0.0071430053,0.0017701054,0.0002836242,0.0011843682,0.006113233,0.012426776,0.066789515,0.12976333,0.055404127,0.71540195],"study_design_scores_gemma":[0.0008455707,0.00079141307,0.01211703,0.0007594335,0.00037656547,0.0025099504,0.0018124164,0.26615694,0.12839088,0.06790546,0.5177606,0.00057370437],"about_ca_topic_score_codex":0.011951623,"about_ca_topic_score_gemma":0.0074897422,"teacher_disagreement_score":0.022036009,"about_ca_system_score_codex":0.001580261,"about_ca_system_score_gemma":0.002449234,"threshold_uncertainty_score":0.07371783},"labels":[],"label_agreement":null},{"id":"W6893400761","doi":"10.5281/zenodo.15571994","title":"Code Complexity Evolution in Student Labs: An Empirical Software Engineering Perspective","year":2025,"lang":"en","type":"preprint","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Cyclomatic complexity; Static program analysis; Search-based software engineering; Software construction; Software; Source lines of code; Bridging (networking); Perspective (graphical); Programming complexity; Empirical research","score_opus":0.06123111379533135,"score_gpt":0.3245316647119341,"score_spread":0.26330055091660276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6893400761","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9951893,0.000121686266,0.0030446283,0.00022066395,0.0000050934214,0.000018996378,0.00074442127,0.000032642867,0.0006226203],"genre_scores_gemma":[0.99558496,0.000063511274,0.0020852736,0.000049993112,0.000019263989,0.0000334528,0.0018686381,0.00002512364,0.00026980284],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99354196,0.0030510179,0.00035089868,0.0011374974,0.0015732923,0.00034525082],"domain_scores_gemma":[0.9129279,0.05170277,0.016673332,0.0071616974,0.008700093,0.0028342218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054457975,0.0002856266,0.00028342984,0.0034932103,0.0005639999,0.0017592233,0.0007088816,0.0006591889,0.0010369434],"category_scores_gemma":[0.05735833,0.0002075947,0.00021863528,0.004808803,0.0013578306,0.0025556157,0.0016391886,0.0012004937,0.00045206974],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012121227,0.0004205189,0.96727383,0.000060299742,0.00006706194,0.0000851861,0.002860547,0.0019265899,0.0013307765,0.0011440694,0.0015833934,0.023126489],"study_design_scores_gemma":[0.000012765311,0.0003046761,0.9761121,0.000036155947,0.00002203644,0.00024999573,0.0031761285,0.01302076,0.0017030182,0.0016640271,0.003664818,0.00003345116],"about_ca_topic_score_codex":0.00223525,"about_ca_topic_score_gemma":0.0030784106,"teacher_disagreement_score":0.0054457975,"about_ca_system_score_codex":0.0007400929,"about_ca_system_score_gemma":0.00064771826,"threshold_uncertainty_score":0.028800428},"labels":[],"label_agreement":null},{"id":"W6893488373","doi":"10.5281/zenodo.16434045","title":"From First Use to Final Commit: Studying the Evolution of Multi-CI Service Adoption (Replication Package)","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University","funders":"","keywords":"Scripting language; Python (programming language); Documentation; Software; Service (business); Software maintenance; Software package; Software evolution","score_opus":0.06518434064067007,"score_gpt":0.27340244919322043,"score_spread":0.20821810855255035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6893488373","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02694996,0.00081098796,0.23157422,0.009805809,0.0028350488,0.004189329,0.34896886,0.22438505,0.1504807],"genre_scores_gemma":[0.104726516,0.0009552651,0.30567363,0.0016834808,0.0004948452,0.010752772,0.3215363,0.13358709,0.120590135],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.99244773,0.0018690936,0.00056124595,0.0010146992,0.0036326498,0.00047448048],"domain_scores_gemma":[0.9059745,0.026889004,0.0037848044,0.023758356,0.03659163,0.003001733],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011819145,0.0013937439,0.0010392809,0.0042124866,0.0016751344,0.005389606,0.0027376555,0.00143415,0.13304201],"category_scores_gemma":[0.08225875,0.0020646427,0.0019674862,0.0056611453,0.0007102242,0.0064801225,0.0049750623,0.0034907914,0.09210244],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032071653,0.00019977578,0.008282431,0.0006502489,0.00009351835,0.000094645715,0.0014362329,0.002134662,0.0015205962,0.008040399,0.85830134,0.11892551],"study_design_scores_gemma":[0.0004401693,0.0006598178,0.060696475,0.0008798418,0.00022501263,0.0002954434,0.0018592016,0.016058646,0.009797748,0.018559996,0.89004385,0.0004839044],"about_ca_topic_score_codex":0.019857358,"about_ca_topic_score_gemma":0.015704343,"teacher_disagreement_score":0.9881809,"about_ca_system_score_codex":0.0023759345,"about_ca_system_score_gemma":0.0066239308,"threshold_uncertainty_score":0.44506985},"labels":[],"label_agreement":null},{"id":"W6894113797","doi":"10.5281/zenodo.7747316","title":"Replication Package - Understanding the Time to First Response In GitHub Pull Requests","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Codebase; Replication (statistics); Exploratory research; Empirical research; Software; Productivity","score_opus":0.05323162045558904,"score_gpt":0.26534230102414846,"score_spread":0.21211068056855942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6894113797","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6631727,0.0012474258,0.2750184,0.0023843383,0.00048410846,0.0014447254,0.00901244,0.023483386,0.023752443],"genre_scores_gemma":[0.9175832,0.0002786514,0.06466294,0.00027951403,0.00014187537,0.0009589702,0.006314385,0.0028723574,0.00690806],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.98861456,0.004294693,0.00090032775,0.0023356457,0.0030191303,0.00083568296],"domain_scores_gemma":[0.844209,0.09098461,0.017472288,0.026183592,0.017737947,0.0034125424],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015046558,0.001194107,0.00090801844,0.003938977,0.0011862862,0.0033304065,0.002299981,0.0015790272,0.006392435],"category_scores_gemma":[0.1238989,0.0011022452,0.0010877822,0.0031296068,0.0013851521,0.0067473347,0.0027119275,0.0022630265,0.0053550964],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019178708,0.0006196988,0.59184355,0.0018207682,0.00045834342,0.0018611125,0.020359047,0.027096555,0.032203965,0.033156686,0.03374793,0.25491446],"study_design_scores_gemma":[0.000234437,0.0013794746,0.53568196,0.00050579995,0.00040751693,0.0024829928,0.010948262,0.29480475,0.014952248,0.038682334,0.09930991,0.0006104173],"about_ca_topic_score_codex":0.0101803,"about_ca_topic_score_gemma":0.005996577,"teacher_disagreement_score":0.98495346,"about_ca_system_score_codex":0.0016508703,"about_ca_system_score_gemma":0.0031069962,"threshold_uncertainty_score":0.079574704},"labels":[],"label_agreement":null},{"id":"W6910986835","doi":"10.5281/zenodo.11125279","title":"Are Large Language Models a Threat to Programming Platforms? An Exploratory Study","year":2024,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Exploratory research; Scripting language; Natural language; Key (lock); Semantics (computer science)","score_opus":0.05163197300563738,"score_gpt":0.2856398672925345,"score_spread":0.2340078942868971,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6910986835","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.979362,0.00025988035,0.0075851292,0.0023602329,0.000017237277,0.00017844938,0.00045364222,0.00015381303,0.009629557],"genre_scores_gemma":[0.9946976,0.00017225342,0.0032742163,0.0002919925,0.000022941442,0.00012175798,0.00037168912,0.00013294992,0.00091471884],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98393166,0.009730983,0.0006727771,0.0011983407,0.0038677803,0.0005984879],"domain_scores_gemma":[0.69783324,0.25958997,0.015162436,0.018878048,0.005680842,0.0028554443],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014163269,0.00046177706,0.0003193567,0.0014059894,0.0015732852,0.0038718618,0.001303343,0.0014596108,0.004985358],"category_scores_gemma":[0.16096106,0.00066313305,0.00050149014,0.0016844943,0.0039191158,0.010417654,0.002857883,0.0040988117,0.0011766623],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005039726,0.008831618,0.5093977,0.0018757178,0.00034944102,0.005287484,0.1365773,0.010821298,0.017851925,0.094210945,0.019166349,0.1905906],"study_design_scores_gemma":[0.00059014274,0.0069262423,0.35899043,0.0024792682,0.0007102605,0.01643602,0.2217763,0.08828024,0.021537112,0.17678507,0.10508433,0.00040455998],"about_ca_topic_score_codex":0.0016174905,"about_ca_topic_score_gemma":0.0016428619,"teacher_disagreement_score":0.014163269,"about_ca_system_score_codex":0.0011167742,"about_ca_system_score_gemma":0.0013063666,"threshold_uncertainty_score":0.07490343},"labels":[],"label_agreement":null},{"id":"W6911815762","doi":"10.5281/zenodo.14041573","title":"Replication Package for `Code Review Comprehension: Reviewing Strategies Seen Through Code Comprehension Theories`","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Code review; Program comprehension; Comprehension; Code (set theory); Static program analysis; Construct (python library); Source code; Software","score_opus":0.055311934500089616,"score_gpt":0.31413262516785173,"score_spread":0.25882069066776214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6911815762","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.111697316,0.0018518,0.4158082,0.021209864,0.03933425,0.23300155,0.05270699,0.01690758,0.10748237],"genre_scores_gemma":[0.24372357,0.0005432997,0.22615296,0.005148965,0.0011279908,0.4877954,0.0094377445,0.0049413447,0.02112868],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.83039516,0.09721868,0.017911576,0.011426735,0.039851084,0.003196668],"domain_scores_gemma":[0.2588164,0.32236165,0.022630382,0.18198697,0.21151446,0.0026902345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.14971785,0.0018058076,0.0023808314,0.006085659,0.0044357227,0.0056031803,0.0059846938,0.0034182018,0.09073239],"category_scores_gemma":[0.70003974,0.0020001421,0.004261507,0.007967032,0.0053224824,0.006270957,0.007011799,0.006835444,0.021740047],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008218601,0.003142557,0.02226143,0.01985599,0.0020678653,0.00110076,0.03614656,0.004857369,0.011132278,0.039701927,0.4552299,0.39628476],"study_design_scores_gemma":[0.010899225,0.005183584,0.13583967,0.026274242,0.001870161,0.00092927454,0.018003186,0.0233429,0.022405973,0.054523066,0.69924456,0.0014842479],"about_ca_topic_score_codex":0.0074270153,"about_ca_topic_score_gemma":0.0059445538,"teacher_disagreement_score":0.14971785,"about_ca_system_score_codex":0.005520588,"about_ca_system_score_gemma":0.013082176,"threshold_uncertainty_score":0.79179317},"labels":[],"label_agreement":null},{"id":"W6912122487","doi":"10.5281/zenodo.16416715","title":"Can LLMs Write CI? A Study on Automatic Generation of GitHub Actions Configurations (Replication Package)","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University","funders":"","keywords":"Documentation; Scripting language; Software; Replication (statistics); Metadata; Rewriting","score_opus":0.06386664889347299,"score_gpt":0.2955125222868887,"score_spread":0.2316458733934157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912122487","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02082512,0.0008657059,0.06567282,0.0051486404,0.0020574322,0.0052634487,0.66231275,0.12799455,0.10985949],"genre_scores_gemma":[0.05846404,0.00076074246,0.11139805,0.002569451,0.00032025523,0.015803367,0.74361736,0.03489481,0.032172002],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935961,0.0023726097,0.0005568701,0.001038018,0.0021692715,0.00026722348],"domain_scores_gemma":[0.9171509,0.04091875,0.0023545083,0.026410503,0.012101115,0.0010642458],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008405761,0.0019508945,0.0013212307,0.0037733964,0.0014256224,0.0035530443,0.004501947,0.0018224568,0.16552448],"category_scores_gemma":[0.0697114,0.0014430487,0.0022527794,0.005161442,0.0008439075,0.005471459,0.0035142743,0.003367035,0.12127208],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003740001,0.00031407122,0.0025667383,0.0011623253,0.00006283572,0.00006460574,0.00033387513,0.0018599413,0.000984478,0.0026870663,0.9307489,0.058841247],"study_design_scores_gemma":[0.001483766,0.0008374967,0.024537502,0.0011904466,0.0002133938,0.00021420713,0.0015579807,0.019737795,0.015262839,0.014105835,0.9205155,0.00034311583],"about_ca_topic_score_codex":0.016061578,"about_ca_topic_score_gemma":0.013802023,"teacher_disagreement_score":0.99159425,"about_ca_system_score_codex":0.0023826112,"about_ca_system_score_gemma":0.0030982518,"threshold_uncertainty_score":0.55373454},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"bench_or_experimental","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"not_applicable","genre":"dataset","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"split"},{"id":"W6912673711","doi":"10.5281/zenodo.8335989","title":"Raw Data for benchmarking structure based domain annotation","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Raw data; Domain (mathematical analysis); Benchmarking; Annotation; Scripting language; Data structure","score_opus":0.04947144449278982,"score_gpt":0.2725789662523704,"score_spread":0.22310752175958057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912673711","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066822637,0.0004967492,0.034960218,0.00031316406,0.00052267406,0.0005520307,0.7165681,0.22428535,0.015619448],"genre_scores_gemma":[0.009332164,0.0002165422,0.025719745,0.00020723323,0.0000342677,0.0008370593,0.9370974,0.023189066,0.0033665402],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9940392,0.0011883929,0.0008316956,0.0015392834,0.0019493648,0.00045195592],"domain_scores_gemma":[0.9877659,0.0033586274,0.0005293469,0.00474931,0.002928479,0.0006683493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051320195,0.0038886615,0.0018741132,0.0058799926,0.002198624,0.004004553,0.0029499035,0.0017884942,0.080457],"category_scores_gemma":[0.025732763,0.0013053444,0.0019880373,0.0064892313,0.00081594655,0.0030487464,0.004413456,0.0036276628,0.123556435],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079319323,0.0003141036,0.0033047276,0.0022023083,0.00017197331,0.0002934832,0.00039570106,0.0034273807,0.008868136,0.003800877,0.93530077,0.041127313],"study_design_scores_gemma":[0.00045560487,0.00024650223,0.013226049,0.0010386687,0.00023345684,0.0007419332,0.0007764911,0.02781747,0.045624398,0.021779003,0.8877586,0.00030181065],"about_ca_topic_score_codex":0.0057685385,"about_ca_topic_score_gemma":0.0061928136,"teacher_disagreement_score":0.080457,"about_ca_system_score_codex":0.0015555618,"about_ca_system_score_gemma":0.003373451,"threshold_uncertainty_score":0.2691555},"labels":[],"label_agreement":null},{"id":"W6912895517","doi":"10.5281/zenodo.7117504","title":"Intro to Git and GitHub","year":2022,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Download; Key (lock); Software; Publication; Control (management)","score_opus":0.022178641717717055,"score_gpt":0.22750950409959114,"score_spread":0.2053308623818741,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912895517","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012559829,0.0065489127,0.20018178,0.005216087,0.005951008,0.0010530351,0.047129568,0.3691265,0.36353707],"genre_scores_gemma":[0.007957686,0.008221909,0.12853295,0.0059465244,0.0038800195,0.0021491582,0.14253122,0.31024715,0.39053348],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971771,0.0004595964,0.0003602834,0.00047956468,0.0011853571,0.00033796465],"domain_scores_gemma":[0.9927302,0.0011069651,0.00032968246,0.002394744,0.0021298479,0.0013085098],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0029836884,0.0027447583,0.0018763159,0.0066374266,0.0015975292,0.008379792,0.00423167,0.002558907,0.44787753],"category_scores_gemma":[0.013225373,0.0018263105,0.0017927882,0.007823188,0.00095203076,0.0069660353,0.007924142,0.005363189,0.6567975],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000677329,0.000044724424,0.000093650306,0.00046007737,0.000010782502,0.00013329396,0.00014409641,0.00025938568,0.0012991579,0.0050428496,0.8919004,0.10054384],"study_design_scores_gemma":[0.00001456851,0.000016444648,0.00031296536,0.00017583984,0.000004730434,0.00025584665,0.00003983491,0.00024594754,0.00044295873,0.004928802,0.9935323,0.000029773602],"about_ca_topic_score_codex":0.0017030953,"about_ca_topic_score_gemma":0.0019472898,"teacher_disagreement_score":0.44787753,"about_ca_system_score_codex":0.0018278046,"about_ca_system_score_gemma":0.0023854033,"threshold_uncertainty_score":0.78753567},"labels":[],"label_agreement":null},{"id":"W6913096877","doi":"10.5281/zenodo.8292979","title":"GPTCloneBench: A comprehensive benchmark of semantic clones and cross-language clones using GPT-3 model and SemanticCloneBench","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Benchmark (surveying); clone (Java method); Ambiguity; Java; Filter (signal processing); Knowledge base; Code (set theory)","score_opus":0.044838139264466344,"score_gpt":0.2885041929786252,"score_spread":0.24366605371415884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6913096877","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5315303,0.0050612427,0.13726038,0.0022975097,0.0014382508,0.0016422245,0.0991841,0.18404017,0.03754591],"genre_scores_gemma":[0.44311315,0.0011444985,0.19193424,0.0013367062,0.00013060223,0.0014714219,0.33515608,0.013919226,0.011794094],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99239916,0.0017676373,0.0008994348,0.0018134859,0.002606078,0.0005141763],"domain_scores_gemma":[0.98291343,0.006698116,0.0009337769,0.003623665,0.0050640763,0.0007669182],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004913734,0.0031333878,0.0009574289,0.0036589287,0.0011259924,0.0022701889,0.0041266284,0.0025001098,0.0036288623],"category_scores_gemma":[0.02291639,0.000684518,0.0020556767,0.004014333,0.0015301687,0.0043019834,0.0027711922,0.0026403817,0.003412323],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023968897,0.0020967352,0.061789848,0.005652988,0.00074320723,0.0025022798,0.0023666269,0.13240662,0.031506646,0.013176657,0.36118737,0.38417414],"study_design_scores_gemma":[0.0007837862,0.002147009,0.032830194,0.00047730742,0.00037320878,0.0017660949,0.0016434927,0.68911654,0.07826836,0.016467236,0.17581865,0.00030810802],"about_ca_topic_score_codex":0.01640988,"about_ca_topic_score_gemma":0.017365107,"teacher_disagreement_score":0.01640988,"about_ca_system_score_codex":0.0023440833,"about_ca_system_score_gemma":0.0028815223,"threshold_uncertainty_score":0.032628715},"labels":[],"label_agreement":null},{"id":"W6913140415","doi":"10.5281/zenodo.6760333","title":"Aligning Programming Language and Natural Language: Exploring Design Choices in Multi-Modal Transformer-Based Embedding for Bug Localization","year":2024,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Embedding; Natural language; Natural (archaeology); First-generation programming language; Programming domain; Replication (statistics); Natural language programming","score_opus":0.056123447161756816,"score_gpt":0.2987255897735646,"score_spread":0.2426021426118078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6913140415","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.050625954,0.00011283067,0.94383776,0.00038132266,0.000024562934,0.00012103706,0.00009891871,0.0026786488,0.0021189454],"genre_scores_gemma":[0.5451978,0.000114180235,0.45117345,0.00016888189,0.000011096959,0.0001739041,0.00025866553,0.0011575853,0.0017445235],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99600804,0.0021929138,0.00022517722,0.00052727415,0.0008187337,0.00022785716],"domain_scores_gemma":[0.9887454,0.0066857045,0.000704144,0.0024558918,0.0012028221,0.00020600166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047735595,0.0006409275,0.0003788289,0.0006057927,0.0004492193,0.0015571091,0.001463126,0.0008055243,0.0039594746],"category_scores_gemma":[0.01847146,0.00060953817,0.00096107594,0.0004883491,0.001957127,0.0053696516,0.003036585,0.0017883214,0.0006434384],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014420167,0.0005864258,0.00746737,0.0018084871,0.00018806626,0.00063038245,0.010546499,0.06683013,0.13020164,0.2936414,0.0054124743,0.4812452],"study_design_scores_gemma":[0.0002091665,0.0010743024,0.0018809543,0.0002918106,0.00038211883,0.00079821626,0.0029980137,0.55524766,0.11067604,0.29795104,0.028337847,0.00015286697],"about_ca_topic_score_codex":0.0009080648,"about_ca_topic_score_gemma":0.0016313787,"teacher_disagreement_score":0.0047735595,"about_ca_system_score_codex":0.0006912812,"about_ca_system_score_gemma":0.0012046341,"threshold_uncertainty_score":0.025245309},"labels":[],"label_agreement":null},{"id":"W6920419628","doi":"10.60692/v4r5x-hjx69","title":"A Refactoring Classification Framework for Efficient Software Maintenance","year":2023,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Code refactoring; Software maintenance; Quality (philosophy); Software quality; Process (computing); Software; Software development","score_opus":0.0647393195551955,"score_gpt":0.267575828737386,"score_spread":0.20283650918219048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920419628","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006158271,0.0008196733,0.98612297,0.00064425124,0.00004797223,0.000831813,0.000319247,0.0008926001,0.0041631944],"genre_scores_gemma":[0.036959212,0.00044786948,0.9599839,0.00008165228,0.000033927125,0.00061763305,0.00078823115,0.00005576625,0.0010318154],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.991951,0.002178796,0.0014208023,0.0014028006,0.0025391688,0.000507437],"domain_scores_gemma":[0.99073035,0.002657692,0.0018496809,0.0010140887,0.0033764136,0.0003718338],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009418496,0.0016561041,0.0012234919,0.012128701,0.0017658628,0.004874159,0.0031149492,0.0018749927,0.0024299475],"category_scores_gemma":[0.012332813,0.00058706437,0.00286752,0.00617406,0.0023173245,0.006298062,0.0023615372,0.002587209,0.0013169171],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010320736,0.00039265168,0.01061538,0.0009830864,0.00013555567,0.0004567118,0.0030692038,0.03371404,0.006271594,0.4440915,0.0076969927,0.49247],"study_design_scores_gemma":[0.00010963216,0.0005382532,0.012926342,0.0016934194,0.00032356265,0.0010637518,0.0018422046,0.5072612,0.0056009907,0.3366633,0.13171019,0.0002670786],"about_ca_topic_score_codex":0.0141751245,"about_ca_topic_score_gemma":0.013777587,"teacher_disagreement_score":0.0141751245,"about_ca_system_score_codex":0.003956529,"about_ca_system_score_gemma":0.006479168,"threshold_uncertainty_score":0.04981041},"labels":[],"label_agreement":null},{"id":"W6923605605","doi":"10.14279/tuj.eceasst.63.915","title":"Investigating Intentional Clone Refactoring","year":2024,"lang":"en","type":"article","venue":"Technische Universität Berlin – Universitätsbibliothek","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Code refactoring; clone (Java method); Commit; Software maintenance; Software evolution; Cloning (programming)","score_opus":0.02041705194927606,"score_gpt":0.2590274233037235,"score_spread":0.23861037135444746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6923605605","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.965151,0.0041118115,0.017499123,0.00049554557,0.00007209905,0.000091482405,0.008589618,0.0009925782,0.0029967031],"genre_scores_gemma":[0.9413188,0.0014444168,0.01863005,0.00030867063,0.00006390386,0.00010055417,0.036080595,0.00024828225,0.001804779],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9921853,0.0015334199,0.00071627315,0.0021519007,0.00281175,0.00060146797],"domain_scores_gemma":[0.93935674,0.029910807,0.010890992,0.011472362,0.0071151904,0.0012539143],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004504844,0.0005836856,0.00063255994,0.0070105735,0.0011197856,0.0015626141,0.0014341193,0.001581884,0.00092452037],"category_scores_gemma":[0.034432728,0.0003868461,0.0007124586,0.0050847414,0.00097118516,0.002629072,0.0018376363,0.0011869757,0.00059072],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024917937,0.00027203976,0.8516338,0.00081859186,0.00029487835,0.001248416,0.0020726891,0.0043009603,0.008636717,0.002549157,0.007884876,0.12003874],"study_design_scores_gemma":[0.00006447442,0.00044698117,0.8320759,0.00034375032,0.00035590187,0.00614634,0.0032416563,0.06291426,0.02270709,0.008207975,0.06335819,0.0001374463],"about_ca_topic_score_codex":0.0077348463,"about_ca_topic_score_gemma":0.012446195,"teacher_disagreement_score":0.0077348463,"about_ca_system_score_codex":0.00084223057,"about_ca_system_score_gemma":0.0012962328,"threshold_uncertainty_score":0.023824155},"labels":[],"label_agreement":null},{"id":"W6929055444","doi":"10.48550/arxiv.2009.02438","title":"A Large Scale Empirical Study of the Impact of Spaghetti Code and Blob Anti-patterns on Program Comprehension","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Program comprehension; Source code; Code (set theory); Empirical research; Task (project management); Comprehension; Scale (ratio); Software maintenance","score_opus":0.10778815700990876,"score_gpt":0.280752169633706,"score_spread":0.1729640126237972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6929055444","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9990287,0.00006337809,0.0003267729,0.000036074052,0.0000022212046,0.0000518659,0.00008313368,0.00000775966,0.0004000804],"genre_scores_gemma":[0.99862146,0.00006888256,0.00077553064,0.000040703508,0.0000049409955,0.00014934743,0.00016749503,0.000011139272,0.0001603886],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9907481,0.0051471554,0.0005425453,0.001125669,0.0020526901,0.00038379634],"domain_scores_gemma":[0.7802908,0.17150867,0.024636671,0.0064664646,0.01390584,0.0031917084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010734453,0.000447046,0.00040271744,0.0013530118,0.00069978356,0.0011073089,0.00069151644,0.00081896776,0.0016651398],"category_scores_gemma":[0.07741518,0.00034521255,0.00038270725,0.0012042164,0.0011301469,0.0014667511,0.001107809,0.0009933572,0.00031450033],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070770475,0.002887409,0.9134076,0.00072702346,0.00032452357,0.00033394777,0.030404197,0.0005053365,0.0049039815,0.00015624211,0.001141663,0.044500362],"study_design_scores_gemma":[0.000050101185,0.001538035,0.98759264,0.00007901892,0.00007998263,0.00018130647,0.006686318,0.00097177416,0.0016734904,0.00008903743,0.0010298245,0.000028424329],"about_ca_topic_score_codex":0.002983409,"about_ca_topic_score_gemma":0.006501334,"teacher_disagreement_score":0.010734453,"about_ca_system_score_codex":0.00078087405,"about_ca_system_score_gemma":0.001086306,"threshold_uncertainty_score":0.056769907},"labels":[],"label_agreement":null},{"id":"W6930219601","doi":"10.5281/zenodo.11939614","title":"Carnet d entretien kia sportage pdf","year":2024,"lang":"fr","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Nucleofection; Work (physics); Context (archaeology)","score_opus":0.02428404758328251,"score_gpt":0.23323725184048347,"score_spread":0.20895320425720096,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6930219601","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00063086586,0.0005925354,0.00054548436,0.00058100634,0.0014744275,0.00009564619,0.0043242206,0.00177079,0.989985],"genre_scores_gemma":[0.0013900261,0.00043559508,0.00035693002,0.00018441583,0.00015754603,0.000027062955,0.001933595,0.000673515,0.9948414],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993691,0.000028225131,0.000027929547,0.00009204905,0.0003946919,0.00008812715],"domain_scores_gemma":[0.99867487,0.00008290792,0.00004035942,0.00012904301,0.0006693078,0.00040358552],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0004780484,0.0012288188,0.0007474967,0.0024438344,0.0017112341,0.00583039,0.0012325282,0.0014719155,0.9291351],"category_scores_gemma":[0.0021149777,0.0005307508,0.0008356415,0.0018059734,0.0006537623,0.0038915959,0.003777134,0.0014784505,0.87053317],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010147731,0.00005213193,0.00015087579,0.00030414565,0.0000051430284,0.00011583845,0.00008032079,0.00005104244,0.0008617069,0.002076587,0.91914713,0.07705373],"study_design_scores_gemma":[0.000007446643,0.000015160874,0.00039556576,0.00006919011,0.0000015431135,0.0000668208,0.000057941474,0.000016090104,0.00020919103,0.0002426235,0.99891317,0.0000051791867],"about_ca_topic_score_codex":0.00576043,"about_ca_topic_score_gemma":0.014142697,"teacher_disagreement_score":0.070864916,"about_ca_system_score_codex":0.0012262902,"about_ca_system_score_gemma":0.0014771881,"threshold_uncertainty_score":0.10108012},"labels":[],"label_agreement":null},{"id":"W6930887008","doi":"10.5281/zenodo.16416714","title":"Can LLMs Write CI? A Study on Automatic Generation of GitHub Actions Configurations (Replication Package)","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University","funders":"","keywords":"Documentation; Scripting language; Software; Replication (statistics); Metadata; Rewriting","score_opus":0.06386664889347299,"score_gpt":0.2955125222868887,"score_spread":0.2316458733934157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6930887008","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02082512,0.0008657059,0.06567282,0.0051486404,0.0020574322,0.0052634487,0.66231275,0.12799455,0.10985949],"genre_scores_gemma":[0.05846404,0.00076074246,0.11139805,0.002569451,0.00032025523,0.015803367,0.74361736,0.03489481,0.032172002],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9935961,0.0023726097,0.0005568701,0.001038018,0.0021692715,0.00026722348],"domain_scores_gemma":[0.9171509,0.04091875,0.0023545083,0.026410503,0.012101115,0.0010642458],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.008405761,0.0019508945,0.0013212307,0.0037733964,0.0014256224,0.0035530443,0.004501947,0.0018224568,0.16552448],"category_scores_gemma":[0.0697114,0.0014430487,0.0022527794,0.005161442,0.0008439075,0.005471459,0.0035142743,0.003367035,0.12127208],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003740001,0.00031407122,0.0025667383,0.0011623253,0.00006283572,0.00006460574,0.00033387513,0.0018599413,0.000984478,0.0026870663,0.9307489,0.058841247],"study_design_scores_gemma":[0.001483766,0.0008374967,0.024537502,0.0011904466,0.0002133938,0.00021420713,0.0015579807,0.019737795,0.015262839,0.014105835,0.9205155,0.00034311583],"about_ca_topic_score_codex":0.016061578,"about_ca_topic_score_gemma":0.013802023,"teacher_disagreement_score":0.99549806,"about_ca_system_score_codex":0.0023826112,"about_ca_system_score_gemma":0.0030982518,"threshold_uncertainty_score":0.55373454},"labels":[],"label_agreement":null},{"id":"W6931780310","doi":"10.5281/zenodo.6625598","title":"Subspecies and Distribution. S.f.floridanusJ.A.Allen,1890—FloridaPeninsulaexceptSEtip(SEUSA). S.f.alacerBangs,1896—CSUSA,inSEKansas,C&amp;SMissouri,EOklahoma,Arkansas,extremeWTennessee,ETexas,mostofMississippi,extremeWCAlabama,andLouisiana.S.fammophilusHowell,1939—EtipoftheFloridaPeninsula(SEUSA).S.faviusOsgood,1910—ConejoI(LosTestigos),offNEVenezuela.S.faztecus].A.Allen,1890—S&amp;EOaxaca,andSWtipofChiapas(SWMexico). S.f.chapmani].A.Allen,1899—W&amp;STexas(SUSA),ECoahuila,N&amp;ENuevoLeon,andmostofTamaulipas(NEMexico). S.f.chiapensisNelson,1904—NEtipofOaxacaandC&amp;SChiapas(Mexico),andWGuatemala. S.f.connectensNelson,1904—STamaulipas,VeracruzexcepttheS,SanLuisPotosi,NEQuerétaro,NEHidalgo,NPuebla,andNOaxaca(EMexico). S.f.continentisOsgood,1912—NWVenezuela. S.f.costaricensisHarris1933—NWCostaRica. S.f.cumanicusThomas,1897—N&amp;CVenezuela. S.f.hesperiusHoffmeister&amp;Lee,1963—NW&amp;CArizona(SWUSA). S.f.hitchensiMearns,1911—SmithandFishermanIsinNorthamptonCounty,EVirginia(EUSA). S.f.holzneriMearns,1896—SEArizona,andSWtipofNewMexico(SWUSA),ESonora,Chihuahua,WDurango,WZacatecas,andESinaloa(NWMexico). S.f.hondurensisGoldman,1932—SEGuatemala,SW&amp;SHonduras,ElSalvador,andNW&amp;CNicaragua. S.fllanensisBlair,1938—SWKansas,SEtipofColorado,WOklahoma,andNCTexas(CUSA). S.f.macrocorpusDiersing&amp;Wilson,1980—Nayarit,SWJalisco,andSWMichoacan(WMexico). S.f.mallurusThomas,1898—EUSA,fromMaineSWtoE&amp;SPennsylvania,andEWestVirginia,WthroughSKentuckytomostofTennessee(excepttheextremeW),andStomostofAlabama(excepttheextremeW)andNFlorida. S.f.margaritaeMiller,1898—MargaritaI,offNVenezuela. S.f.mearnsii].A.Allen,1894—NC&amp;NEUSA,fromMinnesotaEtoMichigan,andStoENebraska,NEKansas,NMissouri,Kentucky,NWtipofVirginia,WWestVirginia,Ohio,N&amp;WPennsylvania,NW&amp;NNewYork,andWVermont,alsoSC&amp;SECanada,inSEOntario,andSEQuebec. S.f.nigronuchalisHartert,1894—ArubaandCuracaoIs(NetherlandsAntilles),offNVenezuela. S.f.orinociThomas,1900—VichadaDepartment(EColombia)andVenezuelaSoftheOrinocoRiver. S.f.orizabaeMerriam,1893—SECoahuila,SWNuevoLeon,SEZacatecas,SanLuisPotosi,Guanajuato,Querétaro,Hidalgo,NWVeracruz,México,andPuebla(CMexico).S.fpaulsoniSchwartz,1956—SEtipoftheFloridaPeninsula(SEUSA). S.f.purgatusThomas,1920—CundinamarcaandTolimaDepartments(CColombia). S.f.restrictusNelson,1907—SWNayarit,C&amp;SJalisco,EColima,andSWMichoacan(Mexico). S.f.russatusJ.A.Allen,1904—SEtipofVeracruz(Mexico). S.f.similisNelson,1907—NCUSA,fromNorthDakota,SWMontana,andNWMinnesotaStoNE&amp;SEWyoming,NEColorado,andNKansas,alsoinSCCanada,SManitoba,andSESaskatchewan(Canada). S.f.subcinctusMiller,1899—Aguascalientes,C&amp;NEJalisco,SWGuanajuato,andNMichoacan(Mexico). S. f. superciliarisJ. A. Allen, 1899 — La Guajira, Magdalena, Atlantico and Bolivar departments (N Colombia). S. [&gt; yucatanicus Miller, 1899 — NW Yucatan Peninsula (SE Mexico). The subspecies mearnsii has been introduced into Vancouver I (SW Canada) and into Washington and Oregon states (NW USA). Also introduced into Italy and the Alps. in Leporidae","year":2016,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Subspecies; Peninsula; Taxonomy (biology); Population; Genetic data","score_opus":0.015393798659136284,"score_gpt":0.22134815005928749,"score_spread":0.2059543514001512,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931780310","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0709876,0.07455039,0.0028320428,0.003732515,0.001277886,0.0007857792,0.035423715,0.0011378133,0.8092723],"genre_scores_gemma":[0.60567987,0.029132217,0.015285829,0.0022652291,0.00072223495,0.00062300405,0.04746575,0.00024825076,0.2985777],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99979347,0.000013218703,0.000016249083,0.000060378083,0.000057532947,0.000059153765],"domain_scores_gemma":[0.9997305,0.000020939304,0.00007522002,0.000012269523,0.00012861937,0.00003243743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00017518426,0.001565496,0.00046475255,0.0036405534,0.001970284,0.00052064675,0.0008078099,0.0005174327,0.041449387],"category_scores_gemma":[0.00035687213,0.0002981969,0.0003418488,0.0028808208,0.0006793319,0.0011097742,0.0009457267,0.00069199543,0.011962568],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026077984,0.00010741422,0.024971168,0.0005794726,0.000049811817,0.00030144333,0.0015051367,0.00035373296,0.0056118337,0.0060595293,0.18969446,0.7705052],"study_design_scores_gemma":[0.000017675975,0.000057932644,0.08583386,0.00019732567,0.000026116573,0.00050932035,0.0004809346,0.00015892684,0.00022849742,0.00063603965,0.91183466,0.000018812281],"about_ca_topic_score_codex":0.18638882,"about_ca_topic_score_gemma":0.27374753,"teacher_disagreement_score":0.18638882,"about_ca_system_score_codex":0.0018346389,"about_ca_system_score_gemma":0.0014321922,"threshold_uncertainty_score":0.37060785},"labels":[],"label_agreement":null},{"id":"W6949165601","doi":"10.5281/zenodo.10828316","title":"Replication Package for \"Effectiveness of ChatGPT for Static Analysis: How Far Are We?\"","year":2024,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Code (set theory); Replication (statistics); Software; Static analysis; Work (physics)","score_opus":0.03584212901976255,"score_gpt":0.2779436582832275,"score_spread":0.24210152926346493,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949165601","genre_codex":"software","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008831625,0.0007256301,0.21984775,0.0034721561,0.0030618159,0.0010249584,0.117799476,0.6166401,0.028596492],"genre_scores_gemma":[0.1060198,0.00092944974,0.34065405,0.0033498795,0.0015035705,0.0073247277,0.1454683,0.33908656,0.05566364],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99168617,0.003330471,0.00070509297,0.0012089779,0.0025425116,0.000526758],"domain_scores_gemma":[0.9187332,0.05214575,0.0025632062,0.01677199,0.00838284,0.0014029494],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009070394,0.0022689255,0.0020194107,0.0029970077,0.001376876,0.0025697716,0.0039635347,0.0022122292,0.16962592],"category_scores_gemma":[0.06431497,0.0016756488,0.0026205298,0.0024828024,0.0011710358,0.0040309615,0.004706929,0.004220884,0.0720488],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021220958,0.00036826165,0.0032918826,0.0019036327,0.00036965118,0.0002338587,0.00071805867,0.0023587826,0.005407093,0.01114581,0.86292875,0.109152175],"study_design_scores_gemma":[0.002951517,0.001590427,0.017679142,0.0014495354,0.000585984,0.0007808899,0.0003781041,0.055264734,0.03346948,0.040935133,0.8442899,0.0006251017],"about_ca_topic_score_codex":0.004064351,"about_ca_topic_score_gemma":0.003413791,"teacher_disagreement_score":0.9909296,"about_ca_system_score_codex":0.0011716924,"about_ca_system_score_gemma":0.002807582,"threshold_uncertainty_score":0.56745523},"labels":[],"label_agreement":null},{"id":"W6949400513","doi":"10.5281/zenodo.15122980","title":"BLAZE : Cross-Language and Cross-Project Bug Localization via Dynamic Chunking and Hard Example Learning","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Waterloo","funders":"","keywords":"Chunking (psychology); Process (computing); R package; Feature (linguistics); Key (lock)","score_opus":0.019450483283077868,"score_gpt":0.28151182347381815,"score_spread":0.2620613401907403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949400513","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003363344,0.00021661445,0.36407948,0.0002962554,0.0003617268,0.00021492825,0.0067091007,0.6207483,0.004010182],"genre_scores_gemma":[0.06055908,0.0003030222,0.6712695,0.0005533411,0.00019168477,0.0008277004,0.03548877,0.21725827,0.0135486545],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99594516,0.00114246,0.00039544894,0.0012239448,0.0009928774,0.00030000883],"domain_scores_gemma":[0.9864991,0.006351219,0.0005418283,0.004662009,0.0014827736,0.00046312687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066414964,0.003803092,0.001696857,0.0028853142,0.0010227174,0.0037590207,0.0044982857,0.0023816582,0.08645321],"category_scores_gemma":[0.026193613,0.002620682,0.0025426464,0.0016999422,0.0012466231,0.0066051167,0.009747158,0.0043865433,0.050732974],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011146519,0.00025398424,0.002744962,0.0014437336,0.00045774726,0.0004223592,0.0010042613,0.007618397,0.012100183,0.010182736,0.51949793,0.4431591],"study_design_scores_gemma":[0.0015105414,0.0007643595,0.0073529715,0.0007725387,0.00037060629,0.0013275194,0.0009061494,0.35654035,0.08787988,0.13123602,0.41065302,0.00068605464],"about_ca_topic_score_codex":0.002199955,"about_ca_topic_score_gemma":0.004114751,"teacher_disagreement_score":0.08645321,"about_ca_system_score_codex":0.00070259545,"about_ca_system_score_gemma":0.0021957962,"threshold_uncertainty_score":0.2892148},"labels":[],"label_agreement":null},{"id":"W6949756894","doi":"10.5281/zenodo.3924700","title":"On the Untriviality of Trivial Packages: An Empirical Study of npm JavaScript Packages","year":2019,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Concordia University; Queen's University","funders":"","keywords":"JavaScript; Unobtrusive JavaScript; Software; Empirical research; Reuse; Popularity; Code (set theory); Simple (philosophy)","score_opus":0.04653909774499646,"score_gpt":0.28316996521608817,"score_spread":0.23663086747109172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949756894","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.996357,0.00017279982,0.0014996708,0.00016493045,0.000006045893,0.000020931571,0.0001792348,0.000028892675,0.0015704116],"genre_scores_gemma":[0.9972953,0.00011952348,0.0015539774,0.00007096155,0.000014065101,0.00003260493,0.00041122324,0.000060916354,0.00044138118],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99255884,0.0034527287,0.0004311219,0.0012667003,0.0018100953,0.00048050564],"domain_scores_gemma":[0.8578404,0.103504926,0.02170451,0.005285066,0.007900701,0.0037643653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009197252,0.00035520268,0.00043744262,0.0037209962,0.0014544812,0.0022193613,0.0012272368,0.0009919315,0.002079457],"category_scores_gemma":[0.07358988,0.0003784176,0.00036919207,0.0037434178,0.0025693316,0.0065027014,0.0020985813,0.0017167601,0.0007787911],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014772265,0.00025864862,0.9412553,0.00028603128,0.00008605195,0.0005170327,0.024798628,0.0004454803,0.001307128,0.0027501918,0.0025343965,0.02561336],"study_design_scores_gemma":[0.000017853532,0.00022229114,0.9339911,0.00020753313,0.00006479376,0.0012834962,0.039466586,0.010461834,0.00092813815,0.0025085614,0.010788181,0.000059612827],"about_ca_topic_score_codex":0.0022392215,"about_ca_topic_score_gemma":0.003329772,"teacher_disagreement_score":0.009197252,"about_ca_system_score_codex":0.0006613132,"about_ca_system_score_gemma":0.00054858014,"threshold_uncertainty_score":0.04864031},"labels":[],"label_agreement":null},{"id":"W6949929472","doi":"10.5281/zenodo.3839074","title":"Characterizing Task-Relevant Information in Natural Language Software Artifacts","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Artifact (error); Natural language; Task (project management); Consistency (knowledge bases); Software; Identification (biology); Set (abstract data type); Software development","score_opus":0.018120083661019887,"score_gpt":0.22331582101500974,"score_spread":0.20519573735398985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949929472","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5025124,0.0018869317,0.32079634,0.0031733462,0.0009137184,0.002677332,0.12293008,0.010361337,0.034748487],"genre_scores_gemma":[0.5224282,0.0006311604,0.31558585,0.0009654378,0.00025817746,0.002611482,0.1471693,0.0013281051,0.009022213],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99732816,0.0011335289,0.00028362905,0.00056200934,0.00058619346,0.00010631324],"domain_scores_gemma":[0.92200005,0.062978104,0.0037783412,0.003493468,0.007019519,0.00073060667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003176958,0.0006242347,0.00029223555,0.0023709375,0.00073287275,0.0019451253,0.00072742696,0.00094262854,0.027658112],"category_scores_gemma":[0.04682869,0.00022820968,0.00039624583,0.0014967688,0.00039514984,0.0020197923,0.0013275936,0.000661513,0.006423256],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020940532,0.0023479408,0.057172272,0.009763328,0.0001970292,0.0020150398,0.009998286,0.006433312,0.11128376,0.02069013,0.27218953,0.5058154],"study_design_scores_gemma":[0.00057876867,0.0019335083,0.23847066,0.0025778662,0.0003608122,0.0044632256,0.010785172,0.13230859,0.122259706,0.05639923,0.42928934,0.00057319493],"about_ca_topic_score_codex":0.0018428338,"about_ca_topic_score_gemma":0.0041285143,"teacher_disagreement_score":0.027658112,"about_ca_system_score_codex":0.0007457485,"about_ca_system_score_gemma":0.0010123001,"threshold_uncertainty_score":0.0925256},"labels":[],"label_agreement":null},{"id":"W6957812177","doi":"10.60692/5ey8q-npz77","title":"MylynSDP — Process - aware artifact filtering based on interest","year":2020,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Artifact (error); Backporting; Function (biology); Software; Task (project management); Software development; Process (computing); Software construction; Software maintenance","score_opus":0.07119640497899826,"score_gpt":0.2465493439908419,"score_spread":0.17535293901184365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6957812177","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14599133,0.00053665106,0.82639927,0.00027597044,0.000062744904,0.00054387556,0.0008107843,0.022339104,0.0030403177],"genre_scores_gemma":[0.6239235,0.00022793912,0.36912492,0.00012731411,0.00004485653,0.0003735762,0.002004244,0.00062199513,0.0035516382],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99792516,0.00038128477,0.00016396861,0.00059916166,0.0008061748,0.00012425371],"domain_scores_gemma":[0.9942009,0.0025250425,0.00083083677,0.0008740222,0.0011567479,0.00041248577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021720538,0.0010209569,0.0007664695,0.0031533602,0.00048585233,0.0016842393,0.0013531722,0.00061713945,0.0018476221],"category_scores_gemma":[0.0074801734,0.00048680682,0.00091377436,0.0014254567,0.00036458953,0.0023830158,0.0024780133,0.0007956408,0.0008990719],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013054322,0.0010548062,0.04703482,0.0008879385,0.00014378158,0.0006456749,0.0016499015,0.0061138533,0.06426409,0.004101419,0.007886558,0.8649118],"study_design_scores_gemma":[0.00035458006,0.0021054987,0.13001154,0.00030018666,0.0005014511,0.0019237448,0.001304755,0.6424374,0.15509333,0.013876877,0.051788364,0.00030228926],"about_ca_topic_score_codex":0.0016651241,"about_ca_topic_score_gemma":0.0019621742,"teacher_disagreement_score":0.0031533602,"about_ca_system_score_codex":0.0005473921,"about_ca_system_score_gemma":0.00075233984,"threshold_uncertainty_score":0.011487067},"labels":[],"label_agreement":null},{"id":"W6967875635","doi":"10.5281/zenodo.16434046","title":"From First Use to Final Commit: Studying the Evolution of Multi-CI Service Adoption (Replication Package)","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University","funders":"","keywords":"Scripting language; Python (programming language); Documentation; Software; Service (business); Software maintenance; Software package; Software evolution","score_opus":0.06518434064067007,"score_gpt":0.27340244919322043,"score_spread":0.20821810855255035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6967875635","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02694996,0.00081098796,0.23157422,0.009805809,0.0028350488,0.004189329,0.34896886,0.22438505,0.1504807],"genre_scores_gemma":[0.104726516,0.0009552651,0.30567363,0.0016834808,0.0004948452,0.010752772,0.3215363,0.13358709,0.120590135],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99244773,0.0018690936,0.00056124595,0.0010146992,0.0036326498,0.00047448048],"domain_scores_gemma":[0.9059745,0.026889004,0.0037848044,0.023758356,0.03659163,0.003001733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011819145,0.0013937439,0.0010392809,0.0042124866,0.0016751344,0.005389606,0.0027376555,0.00143415,0.13304201],"category_scores_gemma":[0.08225875,0.0020646427,0.0019674862,0.0056611453,0.0007102242,0.0064801225,0.0049750623,0.0034907914,0.09210244],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032071653,0.00019977578,0.008282431,0.0006502489,0.00009351835,0.000094645715,0.0014362329,0.002134662,0.0015205962,0.008040399,0.85830134,0.11892551],"study_design_scores_gemma":[0.0004401693,0.0006598178,0.060696475,0.0008798418,0.00022501263,0.0002954434,0.0018592016,0.016058646,0.009797748,0.018559996,0.89004385,0.0004839044],"about_ca_topic_score_codex":0.019857358,"about_ca_topic_score_gemma":0.015704343,"teacher_disagreement_score":0.13304201,"about_ca_system_score_codex":0.0023759345,"about_ca_system_score_gemma":0.0066239308,"threshold_uncertainty_score":0.44506985},"labels":[],"label_agreement":null},{"id":"W6967893722","doi":"10.5281/zenodo.14748996","title":"Replication Package for `Code Review Comprehension: Reviewing Strategies Seen Through Code Comprehension Theories`","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Code review; Program comprehension; Comprehension; Code (set theory); Static program analysis; Construct (python library); Source code; Software","score_opus":0.055311934500089616,"score_gpt":0.31413262516785173,"score_spread":0.25882069066776214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6967893722","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.111697316,0.0018518,0.4158082,0.021209864,0.03933425,0.23300155,0.05270699,0.01690758,0.10748237],"genre_scores_gemma":[0.24372357,0.0005432997,0.22615296,0.005148965,0.0011279908,0.4877954,0.0094377445,0.0049413447,0.02112868],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.83039516,0.09721868,0.017911576,0.011426735,0.039851084,0.003196668],"domain_scores_gemma":[0.2588164,0.32236165,0.022630382,0.18198697,0.21151446,0.0026902345],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.14971785,0.0018058076,0.0023808314,0.006085659,0.0044357227,0.0056031803,0.0059846938,0.0034182018,0.09073239],"category_scores_gemma":[0.70003974,0.0020001421,0.004261507,0.007967032,0.0053224824,0.006270957,0.007011799,0.006835444,0.021740047],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008218601,0.003142557,0.02226143,0.01985599,0.0020678653,0.00110076,0.03614656,0.004857369,0.011132278,0.039701927,0.4552299,0.39628476],"study_design_scores_gemma":[0.010899225,0.005183584,0.13583967,0.026274242,0.001870161,0.00092927454,0.018003186,0.0233429,0.022405973,0.054523066,0.69924456,0.0014842479],"about_ca_topic_score_codex":0.0074270153,"about_ca_topic_score_gemma":0.0059445538,"teacher_disagreement_score":0.85028213,"about_ca_system_score_codex":0.005520588,"about_ca_system_score_gemma":0.013082176,"threshold_uncertainty_score":0.79179317},"labels":[],"label_agreement":null},{"id":"W6968346447","doi":"10.5281/zenodo.3822628","title":"Trials and Tribulations of Developing DDI-Compliant Codebooks at the University of Guelph","year":2002,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Guelph","funders":"","keywords":"Feature (linguistics); Key (lock); Noise (video); Field (mathematics)","score_opus":0.09358244823290503,"score_gpt":0.2531870918166581,"score_spread":0.1596046435837531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6968346447","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7630451,0.0005195434,0.18972114,0.0026803515,0.0007644718,0.012070816,0.0048146173,0.0068774126,0.019506663],"genre_scores_gemma":[0.7007436,0.00021063174,0.26930687,0.0006542754,0.0000444203,0.0070506525,0.005339732,0.0017462274,0.014903571],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9515718,0.035356164,0.0040897285,0.0023999505,0.005640452,0.0009418669],"domain_scores_gemma":[0.73785096,0.1693866,0.0044301436,0.03969086,0.04277218,0.0058692386],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.043986537,0.0006584909,0.0012117985,0.0014857677,0.002207936,0.0024147674,0.0031196466,0.0015605876,0.008811317],"category_scores_gemma":[0.28351834,0.0010319725,0.0005705998,0.0018689121,0.0023479336,0.004438335,0.007492933,0.002378143,0.00411436],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013696073,0.004102934,0.015503463,0.0018402812,0.00015493587,0.00057383033,0.052090034,0.012636939,0.018093737,0.03308615,0.05701323,0.79120845],"study_design_scores_gemma":[0.023252694,0.036661804,0.043509524,0.001814558,0.0007465384,0.0022328205,0.043177243,0.2939527,0.22065221,0.060494673,0.2719683,0.0015369115],"about_ca_topic_score_codex":0.017308285,"about_ca_topic_score_gemma":0.016953357,"teacher_disagreement_score":0.9964231,"about_ca_system_score_codex":0.0035768545,"about_ca_system_score_gemma":0.0059137736,"threshold_uncertainty_score":0.23262584},"labels":[],"label_agreement":null},{"id":"W6968626617","doi":"10.5281/zenodo.3351549","title":"MSRBot: Using Bots to Answer Questions from Software Repositories","year":2019,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Software; Software development; Package development process; Field (mathematics); Process (computing); Work (physics); Software construction; Software analytics; Backporting","score_opus":0.02417584564781291,"score_gpt":0.2477066739102728,"score_spread":0.22353082826245987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6968626617","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38379607,0.0016242096,0.48767206,0.0035561377,0.0007729647,0.0028531698,0.0038082956,0.098125465,0.017791599],"genre_scores_gemma":[0.62380624,0.00053532474,0.34879485,0.0026824337,0.00016486325,0.001870544,0.004416122,0.0021047548,0.015624986],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99467987,0.002706708,0.00036673743,0.00076377555,0.0011192868,0.00036358272],"domain_scores_gemma":[0.971262,0.021256646,0.0021628526,0.002196131,0.0022269115,0.0008955648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005370487,0.0018380188,0.0011384218,0.0024370155,0.0011895797,0.0014009498,0.0018607962,0.0026304666,0.0048398026],"category_scores_gemma":[0.026531257,0.0005883,0.0006548383,0.0008760723,0.0008001316,0.0038950485,0.002747821,0.0016652683,0.003400356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035243009,0.0036118794,0.04550875,0.0055887834,0.00055235275,0.0040256013,0.020080622,0.01094204,0.116159305,0.011384658,0.08558506,0.6930366],"study_design_scores_gemma":[0.00114432,0.004822896,0.04858355,0.0014327915,0.00057893124,0.0047964314,0.010664427,0.6004175,0.0772981,0.056709025,0.19285889,0.0006931279],"about_ca_topic_score_codex":0.0028059098,"about_ca_topic_score_gemma":0.0051034177,"teacher_disagreement_score":0.005370487,"about_ca_system_score_codex":0.0009097145,"about_ca_system_score_gemma":0.0011550044,"threshold_uncertainty_score":0.02840215},"labels":[],"label_agreement":null},{"id":"W6968793479","doi":"10.5281/zenodo.7229976","title":"Ranking Code Clones to Support Maintenance Activities","year":2022,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Code (set theory); Ranking (information retrieval); Replication (statistics); Software maintenance; Key (lock)","score_opus":0.03099441018620133,"score_gpt":0.2473394151671708,"score_spread":0.21634500498096948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6968793479","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09433471,0.0013780717,0.81757385,0.0010595593,0.0009817254,0.0013314557,0.003989878,0.06376469,0.015586063],"genre_scores_gemma":[0.3322397,0.0007615129,0.62950534,0.00036185712,0.0003782166,0.00067793747,0.011750061,0.009185008,0.015140504],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934458,0.0011849004,0.00041810746,0.0007541435,0.0037744704,0.0004226932],"domain_scores_gemma":[0.96413755,0.009360483,0.0028323915,0.0107689295,0.01167617,0.0012244418],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034249313,0.0017504458,0.0013266306,0.004347663,0.0009066469,0.0024285123,0.001685249,0.001488395,0.010784542],"category_scores_gemma":[0.039032463,0.0006775038,0.0012550397,0.0027142689,0.0005282286,0.002847152,0.0016435958,0.0012584901,0.005091331],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010218715,0.00041194446,0.012693237,0.0007147473,0.0001256169,0.00037907777,0.00043321325,0.021212421,0.09011524,0.016816353,0.053760186,0.8023161],"study_design_scores_gemma":[0.0007719795,0.0021723057,0.019933587,0.0005326344,0.000553425,0.002096176,0.0006754076,0.47360685,0.2404346,0.060526013,0.198268,0.00042894934],"about_ca_topic_score_codex":0.0025409998,"about_ca_topic_score_gemma":0.002799051,"teacher_disagreement_score":0.010784542,"about_ca_system_score_codex":0.00086490816,"about_ca_system_score_gemma":0.0019311063,"threshold_uncertainty_score":0.036077857},"labels":[],"label_agreement":null},{"id":"W6969087338","doi":"10.5281/zenodo.7117505","title":"Intro to Git and GitHub","year":2022,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Download; Key (lock); Software; Publication; Control (management)","score_opus":0.022178641717717055,"score_gpt":0.22750950409959114,"score_spread":0.2053308623818741,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6969087338","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012559829,0.0065489127,0.20018178,0.005216087,0.005951008,0.0010530351,0.047129568,0.3691265,0.36353707],"genre_scores_gemma":[0.007957686,0.008221909,0.12853295,0.0059465244,0.0038800195,0.0021491582,0.14253122,0.31024715,0.39053348],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971771,0.0004595964,0.0003602834,0.00047956468,0.0011853571,0.00033796465],"domain_scores_gemma":[0.9927302,0.0011069651,0.00032968246,0.002394744,0.0021298479,0.0013085098],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0029836884,0.0027447583,0.0018763159,0.0066374266,0.0015975292,0.008379792,0.00423167,0.002558907,0.44787753],"category_scores_gemma":[0.013225373,0.0018263105,0.0017927882,0.007823188,0.00095203076,0.0069660353,0.007924142,0.005363189,0.6567975],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000677329,0.000044724424,0.000093650306,0.00046007737,0.000010782502,0.00013329396,0.00014409641,0.00025938568,0.0012991579,0.0050428496,0.8919004,0.10054384],"study_design_scores_gemma":[0.00001456851,0.000016444648,0.00031296536,0.00017583984,0.000004730434,0.00025584665,0.00003983491,0.00024594754,0.00044295873,0.004928802,0.9935323,0.000029773602],"about_ca_topic_score_codex":0.0017030953,"about_ca_topic_score_gemma":0.0019472898,"teacher_disagreement_score":0.44787753,"about_ca_system_score_codex":0.0018278046,"about_ca_system_score_gemma":0.0023854033,"threshold_uncertainty_score":0.78753567},"labels":[],"label_agreement":null},{"id":"W6976524459","doi":"10.60692/m3m93-fq308","title":"Tweaking Association Rules to Optimize Software Change Recommendations","year":2017,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Tweaking; Commit; Association rule learning; Function (biology); Fitness function; Set (abstract data type); Association (psychology); Software","score_opus":0.07681046335864046,"score_gpt":0.2734351798184285,"score_spread":0.19662471645978807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976524459","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30607963,0.0018952597,0.6796268,0.0016803775,0.00038394643,0.0004763317,0.0005564523,0.0051111733,0.0041900068],"genre_scores_gemma":[0.7134724,0.00030359998,0.28200045,0.0004516579,0.0001102273,0.00025401104,0.0008101499,0.00031882338,0.0022786814],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969689,0.0008532557,0.0002351247,0.0009304263,0.0007450047,0.00026723047],"domain_scores_gemma":[0.9856575,0.0099462075,0.0009551146,0.00079562067,0.0023288273,0.00031667136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004820965,0.0016770613,0.0017240274,0.0036415553,0.00082510733,0.001970419,0.002262768,0.002442493,0.0014559951],"category_scores_gemma":[0.028165331,0.00091939594,0.0010382095,0.0024929957,0.0006882705,0.0026110415,0.0009469106,0.0021721923,0.00088685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043662774,0.00078930816,0.042644538,0.00024646363,0.00034684912,0.00039260424,0.00054560514,0.34169334,0.01016824,0.002415983,0.00663069,0.5936898],"study_design_scores_gemma":[0.00005709348,0.00012709989,0.0032723197,0.000033077064,0.000090622816,0.00013424449,0.00013319115,0.9891474,0.0028496855,0.0023298482,0.0017982231,0.000027214859],"about_ca_topic_score_codex":0.010494903,"about_ca_topic_score_gemma":0.015550654,"teacher_disagreement_score":0.010494903,"about_ca_system_score_codex":0.0009390738,"about_ca_system_score_gemma":0.0016044795,"threshold_uncertainty_score":0.025495946},"labels":[],"label_agreement":null},{"id":"W6976596840","doi":"10.60692/7chjc-fxn13","title":"MylynSDP — Process - aware artifact filtering based on interest","year":2020,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Artifact (error); Backporting; Function (biology); Software; Task (project management); Software development; Process (computing); Software construction; Software maintenance","score_opus":0.07119640497899826,"score_gpt":0.2465493439908419,"score_spread":0.17535293901184365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976596840","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14599133,0.00053665106,0.82639927,0.00027597044,0.000062744904,0.00054387556,0.0008107843,0.022339104,0.0030403177],"genre_scores_gemma":[0.6239235,0.00022793912,0.36912492,0.00012731411,0.00004485653,0.0003735762,0.002004244,0.00062199513,0.0035516382],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99792516,0.00038128477,0.00016396861,0.00059916166,0.0008061748,0.00012425371],"domain_scores_gemma":[0.9942009,0.0025250425,0.00083083677,0.0008740222,0.0011567479,0.00041248577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021720538,0.0010209569,0.0007664695,0.0031533602,0.00048585233,0.0016842393,0.0013531722,0.00061713945,0.0018476221],"category_scores_gemma":[0.0074801734,0.00048680682,0.00091377436,0.0014254567,0.00036458953,0.0023830158,0.0024780133,0.0007956408,0.0008990719],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013054322,0.0010548062,0.04703482,0.0008879385,0.00014378158,0.0006456749,0.0016499015,0.0061138533,0.06426409,0.004101419,0.007886558,0.8649118],"study_design_scores_gemma":[0.00035458006,0.0021054987,0.13001154,0.00030018666,0.0005014511,0.0019237448,0.001304755,0.6424374,0.15509333,0.013876877,0.051788364,0.00030228926],"about_ca_topic_score_codex":0.0016651241,"about_ca_topic_score_gemma":0.0019621742,"teacher_disagreement_score":0.0031533602,"about_ca_system_score_codex":0.0005473921,"about_ca_system_score_gemma":0.00075233984,"threshold_uncertainty_score":0.011487067},"labels":[],"label_agreement":null},{"id":"W6976727146","doi":"10.60692/17ja5-za832","title":"Tweaking Association Rules to Optimize Software Change Recommendations","year":2017,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Tweaking; Commit; Association rule learning; Function (biology); Fitness function; Set (abstract data type); Association (psychology); Software","score_opus":0.07681046335864046,"score_gpt":0.2734351798184285,"score_spread":0.19662471645978807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976727146","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30607963,0.0018952597,0.6796268,0.0016803775,0.00038394643,0.0004763317,0.0005564523,0.0051111733,0.0041900068],"genre_scores_gemma":[0.7134724,0.00030359998,0.28200045,0.0004516579,0.0001102273,0.00025401104,0.0008101499,0.00031882338,0.0022786814],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969689,0.0008532557,0.0002351247,0.0009304263,0.0007450047,0.00026723047],"domain_scores_gemma":[0.9856575,0.0099462075,0.0009551146,0.00079562067,0.0023288273,0.00031667136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004820965,0.0016770613,0.0017240274,0.0036415553,0.00082510733,0.001970419,0.002262768,0.002442493,0.0014559951],"category_scores_gemma":[0.028165331,0.00091939594,0.0010382095,0.0024929957,0.0006882705,0.0026110415,0.0009469106,0.0021721923,0.00088685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043662774,0.00078930816,0.042644538,0.00024646363,0.00034684912,0.00039260424,0.00054560514,0.34169334,0.01016824,0.002415983,0.00663069,0.5936898],"study_design_scores_gemma":[0.00005709348,0.00012709989,0.0032723197,0.000033077064,0.000090622816,0.00013424449,0.00013319115,0.9891474,0.0028496855,0.0023298482,0.0017982231,0.000027214859],"about_ca_topic_score_codex":0.010494903,"about_ca_topic_score_gemma":0.015550654,"teacher_disagreement_score":0.010494903,"about_ca_system_score_codex":0.0009390738,"about_ca_system_score_gemma":0.0016044795,"threshold_uncertainty_score":0.025495946},"labels":[],"label_agreement":null},{"id":"W6979299583","doi":"","title":"CODEPROMPTZIP: Code-specific Prompt Compression for Retrieval-Augmented Generation in Coding Tasks with LMs","year":2025,"lang":"en","type":"article","venue":"ArXiv.org","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Code (set theory); Data compression; Coding (social sciences); Focus (optics); Lossy compression; Context (archaeology); Code generation; Encoder; Compression (physics)","score_opus":0.06280035122035256,"score_gpt":0.3020367164221885,"score_spread":0.23923636520183594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6979299583","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09400637,0.0013374506,0.6531565,0.0006559733,0.00039865717,0.0006322322,0.0032741162,0.24214643,0.0043921913],"genre_scores_gemma":[0.35590252,0.0005540679,0.6146848,0.00072465657,0.00014651615,0.0011364074,0.010177704,0.0077299946,0.008943269],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918526,0.00018799295,0.000056868903,0.00025273103,0.00023197035,0.000085283915],"domain_scores_gemma":[0.99670756,0.0016763766,0.00018860973,0.000816371,0.00044665916,0.0001644459],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009806971,0.0019537655,0.0006518948,0.0010425185,0.00039704368,0.00087316224,0.0020071627,0.0010866972,0.007086511],"category_scores_gemma":[0.009469768,0.00044693737,0.000611553,0.00083713187,0.00075278763,0.0026134965,0.002072049,0.0019973246,0.004497739],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014116869,0.00052058836,0.0040365243,0.0011304126,0.00007551403,0.00048135902,0.0007160603,0.027467048,0.08218472,0.004760969,0.052003905,0.8252112],"study_design_scores_gemma":[0.00036675614,0.0010457609,0.002894233,0.000104303734,0.000083304076,0.00059548626,0.0002814924,0.7627737,0.18202513,0.01204368,0.03765796,0.00012822605],"about_ca_topic_score_codex":0.0022261941,"about_ca_topic_score_gemma":0.004121112,"teacher_disagreement_score":0.007086511,"about_ca_system_score_codex":0.0005476808,"about_ca_system_score_gemma":0.0015918631,"threshold_uncertainty_score":0.023706734},"labels":[],"label_agreement":null},{"id":"W6979373137","doi":"","title":"A Study on Mixup-Inspired Augmentation Methods for Software Vulnerability Detection","year":2025,"lang":"en","type":"article","venue":"ArXiv.org","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Software; Metric (unit); Source code; Augment; Code (set theory); Vulnerability (computing); Embedding; Software bug","score_opus":0.07649132692240496,"score_gpt":0.4135226804979202,"score_spread":0.3370313535755152,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6979373137","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16381523,0.006713882,0.81374466,0.0018365396,0.00040710668,0.0002190576,0.0005239047,0.0070766904,0.005662883],"genre_scores_gemma":[0.7372438,0.0013572923,0.25155172,0.0012072354,0.00028857347,0.00030133993,0.0017257057,0.00050930586,0.0058150715],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986682,0.00047433266,0.000073168005,0.0003855606,0.00028779486,0.00011092687],"domain_scores_gemma":[0.9951879,0.0029213422,0.00029635796,0.0009025147,0.0005256651,0.00016618294],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002599475,0.0018655023,0.0012092746,0.0015107364,0.00047323588,0.0012193548,0.0021041678,0.0014939174,0.0018066526],"category_scores_gemma":[0.0100519415,0.00062229665,0.0012815313,0.0011855423,0.0010380929,0.0038689424,0.0022700182,0.0026553771,0.0007312685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007774311,0.00061795523,0.01198146,0.00043160885,0.0003200144,0.00028468168,0.0003142542,0.28118593,0.016998865,0.013335355,0.010262232,0.6634902],"study_design_scores_gemma":[0.000013460536,0.0001501164,0.0005001678,0.000028735889,0.000030232415,0.000093945906,0.000022089842,0.98862445,0.0039032835,0.004415304,0.0022057085,0.000012505674],"about_ca_topic_score_codex":0.0019025918,"about_ca_topic_score_gemma":0.0021603792,"teacher_disagreement_score":0.002599475,"about_ca_system_score_codex":0.00071482663,"about_ca_system_score_gemma":0.0007117833,"threshold_uncertainty_score":0.013747513},"labels":[],"label_agreement":null},{"id":"W6979373727","doi":"","title":"LLM-Based Detection of Tangled Code Changes for Higher-Quality Method-Level Bug Datasets","year":2025,"lang":"en","type":"article","venue":"ArXiv.org","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Winnipeg; University of Manitoba","funders":"","keywords":"Commit; Code (set theory); Granularity; Source code; Classifier (UML); Perceptron; Conflation","score_opus":0.1423767848433097,"score_gpt":0.3971744513948557,"score_spread":0.254797666551546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6979373727","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7768014,0.0039202133,0.103009954,0.0014473292,0.00056379265,0.00041425918,0.03093976,0.07977521,0.0031280948],"genre_scores_gemma":[0.8143834,0.0004280117,0.104332745,0.00044935348,0.000109267276,0.00042318215,0.07472187,0.0012873132,0.0038649077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981647,0.00037108667,0.00016753424,0.00070056855,0.000467959,0.0001281341],"domain_scores_gemma":[0.99303,0.0030448425,0.0008098167,0.0014692324,0.0013124743,0.00033361386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022863126,0.001545949,0.00075887464,0.0028516324,0.0005462416,0.0008411085,0.0017151583,0.001532288,0.0016574315],"category_scores_gemma":[0.015090186,0.00035496554,0.0010708796,0.001500112,0.0006013637,0.00226852,0.0016582183,0.0021886951,0.0015849469],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015121165,0.0017792956,0.13605253,0.0016877004,0.00049224327,0.0014864007,0.0011820425,0.11705633,0.035292964,0.0024260255,0.105458535,0.5955739],"study_design_scores_gemma":[0.00017412452,0.00055685046,0.03021773,0.00010029337,0.000109762834,0.00048183164,0.00034306842,0.92399627,0.026752222,0.0042241826,0.012938472,0.000105177045],"about_ca_topic_score_codex":0.009301102,"about_ca_topic_score_gemma":0.016214589,"teacher_disagreement_score":0.009301102,"about_ca_system_score_codex":0.0010022102,"about_ca_system_score_gemma":0.0013892216,"threshold_uncertainty_score":0.01849389},"labels":[],"label_agreement":null},{"id":"W6986646582","doi":"","title":"Quantifying, Characterizing, and Leveraging Cross-Disciplinary Dependencies: Empirical Studies from a Video Game Development Setting","year":2023,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada)","funders":"","keywords":"Empirical research; Dependency (UML); Process (computing); Code (set theory); Source code; Component (thermodynamics); Key (lock); Software; Graphics","score_opus":0.07503181702512986,"score_gpt":0.32522755486724286,"score_spread":0.250195737842113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6986646582","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99600875,0.00021204982,0.0025680983,0.00008255972,0.0000044291596,0.00009855892,0.00018095702,0.00002060451,0.00082400755],"genre_scores_gemma":[0.9920781,0.00029578994,0.0062065665,0.000073089315,0.000009012832,0.00012834305,0.0007898631,0.000037875954,0.00038132613],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9908296,0.004710883,0.0006442505,0.0012081346,0.0021073774,0.0004997407],"domain_scores_gemma":[0.8340898,0.13215207,0.015172246,0.006189925,0.0094242385,0.0029717532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008817462,0.000590543,0.0004233608,0.0036085215,0.0012066448,0.0018955736,0.0016245657,0.0010559417,0.00095766975],"category_scores_gemma":[0.062071525,0.0005637548,0.0004434752,0.0036781542,0.001966145,0.0039430875,0.0020588767,0.0021317708,0.00033851617],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043955512,0.006327249,0.81520265,0.0012884496,0.0003438024,0.00231821,0.05622792,0.009795519,0.0049015647,0.0033284707,0.003387737,0.09643883],"study_design_scores_gemma":[0.00007810228,0.0012093569,0.88147897,0.0005429455,0.00021431388,0.001077586,0.061055884,0.035038445,0.004239239,0.0026916717,0.012230894,0.00014265168],"about_ca_topic_score_codex":0.00955401,"about_ca_topic_score_gemma":0.018319279,"teacher_disagreement_score":0.00955401,"about_ca_system_score_codex":0.0014616256,"about_ca_system_score_gemma":0.0015934128,"threshold_uncertainty_score":0.046631753},"labels":[],"label_agreement":null},{"id":"W6997367243","doi":"","title":"Using Repository Level Embeddings to Generate Code Using Large Language Models","year":2024,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Code (set theory); Source code; Natural language; Key (lock); Set (abstract data type)","score_opus":0.05795592383751554,"score_gpt":0.3158428256601142,"score_spread":0.25788690182259866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6997367243","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11952125,0.00027719323,0.84344053,0.00062948046,0.00029467276,0.00021712104,0.001985858,0.029934216,0.003699607],"genre_scores_gemma":[0.5127759,0.00018536767,0.46715587,0.00020652039,0.000073855765,0.0002989542,0.009784452,0.0049233357,0.004595642],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921143,0.00026543878,0.00005186876,0.00022107623,0.00018515332,0.00006510355],"domain_scores_gemma":[0.9959086,0.0022330235,0.00021316063,0.00075486937,0.00078435184,0.000106050094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007772586,0.0009961046,0.00051312556,0.0013321425,0.00048774088,0.0010823342,0.0010470599,0.0010480724,0.003765988],"category_scores_gemma":[0.0070890747,0.0007287199,0.0012303554,0.0009165835,0.0004837665,0.0029984873,0.0014941493,0.0017898304,0.0033523396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005379091,0.0005671832,0.00950924,0.00044965852,0.000142789,0.00067463703,0.00063696277,0.25062302,0.023221426,0.022929514,0.034930736,0.65577704],"study_design_scores_gemma":[0.00002639686,0.00005148876,0.0002677125,0.000015466292,0.00002329257,0.000059926788,0.00006698891,0.9804782,0.0053480132,0.011381308,0.002267946,0.000013233159],"about_ca_topic_score_codex":0.0022579364,"about_ca_topic_score_gemma":0.005201836,"teacher_disagreement_score":0.003765988,"about_ca_system_score_codex":0.00065888424,"about_ca_system_score_gemma":0.0010051894,"threshold_uncertainty_score":0.012598515},"labels":[],"label_agreement":null},{"id":"W6999429681","doi":"","title":"A Comprehensive Study on Quality Aspects and Industry Perspective in Backporting","year":2024,"lang":"en","type":"dissertation","venue":"University Library (University of Saskatchewan)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software quality analyst; Software quality; Software quality control; Quality (philosophy); Software quality assurance; Quality assurance; Software quality management; Software maintenance; Software; Technical debt","score_opus":0.019571369115488617,"score_gpt":0.25166576491723136,"score_spread":0.23209439580174274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6999429681","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6819756,0.030999042,0.10396916,0.012358417,0.00046349125,0.0003551522,0.00038649546,0.00049803685,0.16899459],"genre_scores_gemma":[0.9530055,0.011856011,0.024278888,0.0010575495,0.00016470946,0.000101058125,0.0003435052,0.00019366047,0.008999095],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99122345,0.0017854982,0.0005526809,0.0010055053,0.0047438736,0.0006889777],"domain_scores_gemma":[0.94418925,0.025420725,0.010796891,0.004349203,0.013503734,0.0017401665],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009254954,0.00047877576,0.00023264003,0.0052225925,0.0014977063,0.005350577,0.0009728627,0.00108822,0.0036172194],"category_scores_gemma":[0.034440346,0.00038706354,0.0006018976,0.008432455,0.0031198775,0.009512871,0.002611594,0.0021383741,0.0005449893],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012882285,0.00031263838,0.13334575,0.0021754121,0.000097182274,0.0011290346,0.04062366,0.0044414965,0.0065847468,0.17478165,0.0100964615,0.6262831],"study_design_scores_gemma":[0.000030379422,0.00081068644,0.38164446,0.0052894508,0.0002647842,0.0030287032,0.054389127,0.014326522,0.0088785775,0.06407417,0.4670276,0.00023555181],"about_ca_topic_score_codex":0.0045414916,"about_ca_topic_score_gemma":0.004130697,"teacher_disagreement_score":0.009254954,"about_ca_system_score_codex":0.004188751,"about_ca_system_score_gemma":0.0034154626,"threshold_uncertainty_score":0.048945427},"labels":[],"label_agreement":null},{"id":"W7006846287","doi":"","title":"Zero-shot synthesis of compilable code for incomplete code snippets using LLMs","year":2024,"lang":"en","type":"dissertation","venue":"Mspace (University of Manitoba)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Python (programming language); Validator; Code reuse; Benchmark (surveying); Code (set theory); Software","score_opus":0.046215664354894714,"score_gpt":0.26485987567339137,"score_spread":0.21864421131849665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7006846287","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10726914,0.0024150307,0.5076014,0.00088216964,0.00081669097,0.00059273763,0.018857688,0.3534655,0.008099696],"genre_scores_gemma":[0.2534531,0.00087459985,0.63040507,0.0011821035,0.00022140292,0.0010287977,0.071103476,0.031597342,0.010134135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977114,0.0003572723,0.0001657457,0.00073514297,0.00084999617,0.00018026667],"domain_scores_gemma":[0.995033,0.0024284422,0.00039157952,0.0009800767,0.0010231367,0.00014374583],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013875036,0.0031705857,0.0011420964,0.0035687606,0.00084869633,0.0015057118,0.0021946388,0.0014961767,0.0059892936],"category_scores_gemma":[0.010098506,0.00088686525,0.002175435,0.0015099283,0.0012871989,0.0019088837,0.0023708944,0.0015767367,0.004855917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018817344,0.00041660757,0.018852262,0.0036362726,0.00045405905,0.0022751414,0.0013250385,0.04746292,0.09494926,0.006502517,0.12819867,0.69404554],"study_design_scores_gemma":[0.0003116081,0.00059819146,0.010876405,0.00042731414,0.00031032405,0.0013911008,0.0007762952,0.6960389,0.1732063,0.02312976,0.09267896,0.0002548286],"about_ca_topic_score_codex":0.004150822,"about_ca_topic_score_gemma":0.011724036,"teacher_disagreement_score":0.0059892936,"about_ca_system_score_codex":0.0008345255,"about_ca_system_score_gemma":0.002064578,"threshold_uncertainty_score":0.020036161},"labels":[],"label_agreement":null},{"id":"W7014864557","doi":"","title":"Practical Considerations of Deploying AI in Defect Prediction: A Case Study within the Turkish Telecommunication Industry","year":2009,"lang":"en","type":"article","venue":"NPARC","topic":"Software Engineering Research","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Türkiye Bilimsel ve Teknolojik Araştırma Kurumu","keywords":"Turkish; Key (lock); Field (mathematics); Process (computing)","score_opus":0.05424820902952547,"score_gpt":0.3502673163149303,"score_spread":0.2960191072854048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7014864557","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9589113,0.00022186435,0.018139131,0.0030306946,0.00006726602,0.0003229268,0.00005409118,0.00031615395,0.018936694],"genre_scores_gemma":[0.9884595,0.00009537675,0.009348511,0.00008835448,0.000012183089,0.000039446495,0.000030501664,0.000028282479,0.0018979126],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9954242,0.0030399829,0.0001718476,0.00024347141,0.00072940666,0.0003910739],"domain_scores_gemma":[0.9744664,0.018369274,0.0010851105,0.0012833296,0.0037230612,0.0010728583],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005202098,0.0005125551,0.0002830256,0.0011426384,0.0021889075,0.0024520692,0.0018888282,0.0021504825,0.0044683567],"category_scores_gemma":[0.019829651,0.00027440523,0.00027060093,0.0010333394,0.001058817,0.0020111373,0.0011732085,0.0010109677,0.0009367518],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016225234,0.005608109,0.18330024,0.0008704481,0.000118542535,0.03171602,0.03761487,0.06377004,0.032338828,0.013397933,0.011603678,0.61803865],"study_design_scores_gemma":[0.0005797037,0.009671298,0.15895791,0.00083093304,0.00034787896,0.021105241,0.18627743,0.50158304,0.029861802,0.021100255,0.06934532,0.0003392524],"about_ca_topic_score_codex":0.00795646,"about_ca_topic_score_gemma":0.010826121,"teacher_disagreement_score":0.00795646,"about_ca_system_score_codex":0.001229592,"about_ca_system_score_gemma":0.0023460735,"threshold_uncertainty_score":0.027511656},"labels":[],"label_agreement":null},{"id":"W7015086753","doi":"","title":"Security Evaluations of GitHub's Copilot","year":2023,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada)","funders":"","keywords":"Vulnerability (computing); Set (abstract data type); Code (set theory); Vulnerability management; Vulnerability assessment","score_opus":0.019828913603516518,"score_gpt":0.25794415196958526,"score_spread":0.23811523836606874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015086753","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9324548,0.0022743775,0.014487269,0.0012560544,0.00038676822,0.0010072266,0.0038008282,0.028304243,0.016028408],"genre_scores_gemma":[0.87958425,0.0014665623,0.077464215,0.0010232378,0.000111491405,0.0010488126,0.023730347,0.007098513,0.0084726475],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99423075,0.0020923726,0.0004129388,0.0006535392,0.0021583287,0.00045213592],"domain_scores_gemma":[0.969218,0.01755693,0.0022277005,0.0046874704,0.004793236,0.0015166374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056824116,0.0014247457,0.00069214014,0.0027351258,0.0006452734,0.0012669602,0.0027357396,0.0015400486,0.0031970749],"category_scores_gemma":[0.032293182,0.00053745025,0.0008563615,0.0012665482,0.0011421051,0.0019235829,0.0024160852,0.0019910813,0.0018214101],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013416557,0.007024542,0.06318218,0.009178644,0.0010753179,0.0039320663,0.012570995,0.04535922,0.069695674,0.007974947,0.22930251,0.53728735],"study_design_scores_gemma":[0.003890619,0.017339546,0.16410607,0.0026268768,0.00061353546,0.006339153,0.004966899,0.43477735,0.12038838,0.0066710943,0.23750907,0.0007713527],"about_ca_topic_score_codex":0.0033978766,"about_ca_topic_score_gemma":0.0058418163,"teacher_disagreement_score":0.0056824116,"about_ca_system_score_codex":0.0013341702,"about_ca_system_score_gemma":0.0011647652,"threshold_uncertainty_score":0.030051827},"labels":[],"label_agreement":null},{"id":"W7015636045","doi":"","title":"Trade-Off Exploration for Acceleration of Continuous Integration","year":2023,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada)","funders":"","keywords":"Acceleration; Trustworthiness; Duration (music); Product (mathematics); Process (computing); Spectral acceleration","score_opus":0.026695951121450623,"score_gpt":0.24477029514447934,"score_spread":0.21807434402302872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015636045","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58147544,0.0047903233,0.36706793,0.0032780452,0.0004068248,0.0006746454,0.00031170784,0.016690932,0.025304126],"genre_scores_gemma":[0.83523,0.00039598384,0.15991907,0.0003003911,0.00007140224,0.00025758927,0.00035796544,0.0014342276,0.0020333729],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98234516,0.006141914,0.0011200126,0.0020148715,0.0070520146,0.0013259546],"domain_scores_gemma":[0.8755494,0.07223721,0.009790763,0.028890077,0.010868281,0.0026642636],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01559786,0.0018712227,0.0008111127,0.002130187,0.00078735896,0.003910828,0.0025127467,0.0014184545,0.0054833973],"category_scores_gemma":[0.11254636,0.0011688059,0.0010745868,0.0015772014,0.0019265108,0.0072466405,0.0043458696,0.003263273,0.0018778347],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004354557,0.0009254699,0.04135608,0.0013753391,0.00033778706,0.0007383832,0.0029850074,0.094778106,0.080264315,0.03317336,0.009697797,0.7300138],"study_design_scores_gemma":[0.0007176018,0.007769825,0.042735446,0.00068425405,0.0006472369,0.0015656773,0.0023157024,0.7302178,0.0740128,0.07885733,0.06017356,0.0003027719],"about_ca_topic_score_codex":0.0013982833,"about_ca_topic_score_gemma":0.0014541951,"teacher_disagreement_score":0.01559786,"about_ca_system_score_codex":0.0013734826,"about_ca_system_score_gemma":0.002217721,"threshold_uncertainty_score":0.082490385},"labels":[],"label_agreement":null},{"id":"W7015853340","doi":"","title":"Variability in Factors Influencing Pull Request Merge Decisions: A Microscopic Exploration","year":2024,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada)","funders":"","keywords":"Merge (version control); Logistic regression; Regression analysis; Regression","score_opus":0.019893854708766204,"score_gpt":0.24560292378514487,"score_spread":0.22570906907637867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015853340","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9903832,0.00012607174,0.0066814944,0.00041280207,0.000006263372,0.000089594636,0.00013158412,0.00004251727,0.0021265247],"genre_scores_gemma":[0.9988557,0.000030961164,0.00086696376,0.00002245487,0.0000031031148,0.000023918683,0.000041682077,0.000009831067,0.00014542662],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9896344,0.006104439,0.0004412798,0.0014563365,0.0016210633,0.00074256404],"domain_scores_gemma":[0.8635353,0.11374352,0.012935855,0.0046950984,0.0037210754,0.001369163],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013726218,0.0003415566,0.00049354805,0.0015701402,0.0009835871,0.0032222103,0.00097481906,0.00082558946,0.0032110019],"category_scores_gemma":[0.07601593,0.00038958603,0.00066709414,0.0016845849,0.0020930734,0.003191556,0.0023366094,0.0013379469,0.0003642536],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030025924,0.00037158147,0.923913,0.0002864613,0.0002866359,0.0005775609,0.014122742,0.010864185,0.003726605,0.007184953,0.0010701485,0.03729592],"study_design_scores_gemma":[0.00003761491,0.0004646822,0.9168979,0.00010401035,0.00019127694,0.00035346992,0.015995167,0.051872615,0.00159005,0.0092393095,0.003157563,0.00009646835],"about_ca_topic_score_codex":0.0075484086,"about_ca_topic_score_gemma":0.008000398,"teacher_disagreement_score":0.013726218,"about_ca_system_score_codex":0.0017139056,"about_ca_system_score_gemma":0.0018337257,"threshold_uncertainty_score":0.07259208},"labels":[],"label_agreement":null},{"id":"W7023414520","doi":"","title":"Software Journeys","year":2013,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Documentation; Software development; Software; Software construction; Software documentation; Backporting; Software system; Code (set theory); Source code; Resource (disambiguation)","score_opus":0.01091135535787974,"score_gpt":0.20949947220232254,"score_spread":0.1985881168444428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7023414520","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11284724,0.0028872578,0.061863426,0.016291842,0.0012994297,0.00040750153,0.0028879559,0.0031068414,0.79840845],"genre_scores_gemma":[0.42333454,0.004624455,0.07280321,0.0018158637,0.00016005168,0.0003552514,0.004600582,0.0025462133,0.48975986],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99815303,0.00049665757,0.00011260141,0.00034485632,0.0006694251,0.00022347191],"domain_scores_gemma":[0.99624556,0.00064361584,0.00032198167,0.0006936995,0.0010793129,0.0010158563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010682513,0.0005543935,0.00024870082,0.0018335616,0.0037488723,0.0063001034,0.0010329762,0.001015006,0.06072701],"category_scores_gemma":[0.009508266,0.00043275996,0.0005743779,0.0017057745,0.0022766718,0.006398323,0.0065759486,0.0020131576,0.022295628],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014099402,0.00013733676,0.0110698985,0.0005006123,0.00003298554,0.0010135288,0.05886668,0.0012751671,0.0038367964,0.41636205,0.13841175,0.3683521],"study_design_scores_gemma":[0.0000062337804,0.000053864074,0.0029832493,0.00015387089,0.0000056196795,0.00047038088,0.010652341,0.0002860993,0.00039030303,0.030033078,0.95493996,0.000025017418],"about_ca_topic_score_codex":0.006042431,"about_ca_topic_score_gemma":0.008526051,"teacher_disagreement_score":0.06072701,"about_ca_system_score_codex":0.0022640184,"about_ca_system_score_gemma":0.003826088,"threshold_uncertainty_score":0.20315212},"labels":[],"label_agreement":null},{"id":"W7024070025","doi":"","title":"Queerly Platonic: Constellating An Asian North American Critique of Compulsory Sexuality","year":2024,"lang":"en","type":"dissertation","venue":"MacSphere (McMaster University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Human sexuality; Queer; Asexuality; Natural (archaeology); Lesbian; Polynesians","score_opus":0.014016378996288595,"score_gpt":0.2572893968503842,"score_spread":0.2432730178540956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024070025","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36433175,0.0040128333,0.0044211927,0.03944546,0.00053960143,0.000035224708,0.000034096607,0.000053966553,0.5871259],"genre_scores_gemma":[0.9628646,0.0012894174,0.0005709047,0.0015096613,0.00006757465,0.000023518902,0.0000102910635,0.00003181152,0.033632293],"study_design_codex":"qualitative","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99855596,0.0007926495,0.000019978506,0.0001401695,0.00020930078,0.00028190055],"domain_scores_gemma":[0.99872833,0.00064565835,0.000082947394,0.00011017717,0.0002483373,0.00018452224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027616906,0.00042034555,0.00023097334,0.0009530745,0.01773001,0.0071550575,0.0009329013,0.0012528576,0.0034856894],"category_scores_gemma":[0.0019670818,0.00018425587,0.00016306913,0.0010363068,0.035605013,0.004985654,0.0032413555,0.003637018,0.00028095598],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000109727725,0.000011710592,0.0006311589,0.000023431428,0.0000015239439,0.00018014885,0.5645524,0.000041903102,0.00031737526,0.42583367,0.00315582,0.0052397894],"study_design_scores_gemma":[0.000005024885,0.00001976731,0.002164879,0.00015734427,0.000008925695,0.0002023385,0.6461987,0.0003015298,0.0005768,0.040930316,0.3094094,0.0000249165],"about_ca_topic_score_codex":0.25340828,"about_ca_topic_score_gemma":0.3942327,"teacher_disagreement_score":0.25340828,"about_ca_system_score_codex":0.013166026,"about_ca_system_score_gemma":0.010939236,"threshold_uncertainty_score":0.50386655},"labels":[],"label_agreement":null},{"id":"W7024215644","doi":"","title":"Regulation of the human oxytocin receptor gene by interleukin-1Ã and interleukin-6 in vitro","year":2000,"lang":"en","type":"other","venue":"Library and Archives Canada (Government of Canada)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"In vitro; Oxytocin; Gene; Receptor; Oxytocin receptor; Gene expression; Human placenta; Signal transduction","score_opus":0.0025912217478481535,"score_gpt":0.1485249513329135,"score_spread":0.14593372958506534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024215644","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97976744,0.005645753,0.0036628356,0.00047603023,0.00024882768,0.00008147677,0.0016532218,0.00006850029,0.008395925],"genre_scores_gemma":[0.97773176,0.0021322367,0.004120255,0.00021088833,0.00008293137,0.000071688075,0.004302717,0.000076391945,0.011271064],"study_design_codex":"bench_or_experimental","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994886,0.00017480229,0.000052905136,0.00007218239,0.00009943269,0.00011201128],"domain_scores_gemma":[0.9988116,0.0008718391,0.000056228433,0.00010489848,0.000089895475,0.000065538894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000471737,0.00025232983,0.00040116502,0.0005608856,0.00034467372,0.00078045076,0.00033954982,0.00025648918,0.0036546718],"category_scores_gemma":[0.0006146134,0.000319893,0.00062454556,0.0003250492,0.0004744294,0.00026142003,0.00033304974,0.0009043477,0.0014635266],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007226443,0.00014243349,0.00040031274,0.000060046023,0.000009349754,0.000121202174,0.00016451112,0.00017227583,0.9958228,0.0004349736,0.00026426854,0.0016851986],"study_design_scores_gemma":[0.00014142318,0.00087549543,0.011651428,0.000046150963,0.000040964536,0.00041743758,0.00034290928,0.003113606,0.97508097,0.0003458211,0.007921471,0.000022405222],"about_ca_topic_score_codex":0.0033436012,"about_ca_topic_score_gemma":0.003932039,"teacher_disagreement_score":0.0036546718,"about_ca_system_score_codex":0.0005955898,"about_ca_system_score_gemma":0.0004004946,"threshold_uncertainty_score":0.012226045},"labels":[],"label_agreement":null},{"id":"W7028187105","doi":"","title":"An Empirical Analysis of Java Language Use in Open Source Applications","year":2019,"lang":"en","type":"dissertation","venue":"QSpace (Queen's University Library)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Nucleofection; Gestational period; Demotion; Hyporeflexia; Subpoena; Pretext; Hemopericardium; TSG101","score_opus":0.012511632806941318,"score_gpt":0.26559344998672724,"score_spread":0.2530818171797859,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7028187105","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99672526,0.00013605309,0.0004215295,0.00010319602,0.0000043944324,0.000017481521,0.0018669829,0.000030140442,0.00069493795],"genre_scores_gemma":[0.9913585,0.00015539658,0.0012048266,0.00007296125,0.000013649577,0.000073492134,0.006619309,0.000064804764,0.00043707108],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9943982,0.001861523,0.0006492036,0.0010583868,0.0016526314,0.00038010193],"domain_scores_gemma":[0.8887286,0.07646257,0.017753053,0.005587336,0.009565166,0.0019031758],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0040942905,0.00024664073,0.00028734843,0.003580299,0.00052257685,0.0016786715,0.0007175113,0.00076577737,0.00097389123],"category_scores_gemma":[0.05283199,0.00028088392,0.00035996322,0.0049400805,0.00085334363,0.002525158,0.0012279155,0.0012147282,0.0006521512],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013286533,0.00022425398,0.97919077,0.00014227438,0.00007334736,0.00008972353,0.0030731808,0.00021071776,0.0009668713,0.00018981253,0.0018027229,0.013903296],"study_design_scores_gemma":[0.000006544382,0.00006288659,0.9920528,0.00004444318,0.000023379318,0.00020430844,0.0024509262,0.0016865209,0.00055678014,0.00013997043,0.0027541304,0.000017339242],"about_ca_topic_score_codex":0.0028099315,"about_ca_topic_score_gemma":0.0042147995,"teacher_disagreement_score":0.9959057,"about_ca_system_score_codex":0.00055952853,"about_ca_system_score_gemma":0.00047917754,"threshold_uncertainty_score":0.021652937},"labels":[],"label_agreement":null},{"id":"W7030920646","doi":"","title":"Changes of State","year":2009,"lang":"en","type":"other","venue":"Bulletin of Miscellaneous Information (Royal Gardens Kew)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.006776409952344637,"score_gpt":0.19542605108867908,"score_spread":0.18864964113633445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7030920646","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01121832,0.001026621,0.0010880579,0.0046415813,0.0028964123,0.00015755833,0.017201414,0.000975394,0.9607947],"genre_scores_gemma":[0.05463388,0.0009761433,0.0007760766,0.0014180524,0.0002159338,0.00009450892,0.01009441,0.00055404456,0.931237],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984041,0.00007829301,0.00005706616,0.00027141813,0.00079789513,0.00039134934],"domain_scores_gemma":[0.99679655,0.00014114469,0.00009610655,0.0004838968,0.0017877199,0.0006945278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010689711,0.00023891771,0.0003544296,0.0013542199,0.0034948776,0.0039638607,0.0010361228,0.0009927384,0.14907445],"category_scores_gemma":[0.0057352255,0.00025741966,0.00035110858,0.0021380954,0.00091012067,0.0017400242,0.0021573121,0.0012027909,0.03804872],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017328453,0.000029298013,0.007937855,0.00019780149,0.000014718133,0.00025923562,0.004224788,0.00009790119,0.0014738747,0.03551791,0.8304268,0.11964646],"study_design_scores_gemma":[0.000003877735,0.0000056409895,0.007399009,0.00003888046,0.000003417559,0.000025191437,0.000579622,0.000021325366,0.0000915255,0.00035505486,0.9914689,0.000007663299],"about_ca_topic_score_codex":0.50977933,"about_ca_topic_score_gemma":0.6971033,"teacher_disagreement_score":0.50977933,"about_ca_system_score_codex":0.009819604,"about_ca_system_score_gemma":0.00854926,"threshold_uncertainty_score":0.98621535},"labels":[],"label_agreement":null},{"id":"W7034512950","doi":"","title":"Using shore-based surveys to assess vessel traffic patterns in two migratory bird sanctuaries","year":2022,"lang":"en","type":"article","venue":"Western CEDAR (Western Washington University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Harbour; Habitat; Overwintering; Automatic Identification System; Shoal; Baseline (sea)","score_opus":0.065069998999181,"score_gpt":0.29247806177855024,"score_spread":0.22740806277936923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7034512950","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99953103,0.000009896808,0.00015481538,0.00000398721,0.0000015973142,0.000018585031,0.00008913093,0.0000028536044,0.0001882712],"genre_scores_gemma":[0.9971547,0.000039270293,0.0017758376,0.000016659516,0.000005818212,0.000056135803,0.000680991,0.0000030996125,0.00026750876],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9995454,0.00012263887,0.00003693161,0.00012892648,0.000087507164,0.0000785858],"domain_scores_gemma":[0.99861,0.00014246783,0.00052727177,0.0000713594,0.00036103945,0.00028783848],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004769151,0.0002987438,0.00016465956,0.0009291884,0.0004576046,0.0003405058,0.0002854243,0.0003072538,0.0005366049],"category_scores_gemma":[0.0013867994,0.00022877786,0.00017077391,0.000574356,0.00031622974,0.0004685826,0.00052417425,0.00024137936,0.00018522415],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000063875756,0.000071843686,0.993082,0.000010797657,0.000018845802,0.00006196119,0.0005861427,0.00007331421,0.0018558722,0.000012227672,0.0000639239,0.0040993076],"study_design_scores_gemma":[0.0000037509094,0.0002310765,0.9987124,0.0000028242694,0.0000055563355,0.00004911295,0.00046623408,0.00027000124,0.000115813986,0.0000055071314,0.00013444628,0.000003359944],"about_ca_topic_score_codex":0.02141185,"about_ca_topic_score_gemma":0.08918896,"teacher_disagreement_score":0.02141185,"about_ca_system_score_codex":0.00044970986,"about_ca_system_score_gemma":0.0004071617,"threshold_uncertainty_score":0.042574406},"labels":[],"label_agreement":null},{"id":"W7034651932","doi":"","title":"Two-dimensional waveform tomography of the Queen Charlotte Basin of Western Canada and the Seattle fault zone","year":2011,"lang":"en","type":"dissertation","venue":"Summit (Simon Fraser University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Tomography; Waveform; Inversion (geology); Attenuation; Structural basin; Reflection (computer programming); Seismic tomography; Fault (geology)","score_opus":0.00683281761066355,"score_gpt":0.19198092123494126,"score_spread":0.18514810362427772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7034651932","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99611485,0.000029292167,0.0018017775,0.000044714747,0.00000287056,0.000020772248,0.0004850031,0.00008322294,0.0014174937],"genre_scores_gemma":[0.99590224,0.000047833328,0.0026533327,0.00000848314,0.0000011023325,0.00001012127,0.0005419304,0.0000128745205,0.0008220997],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99992156,0.000003991222,0.000003870547,0.000016199638,0.000035926038,0.00001857588],"domain_scores_gemma":[0.999816,0.000024228355,0.000024180574,0.0000111176805,0.00009097637,0.000033554916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00008566801,0.00034421828,0.00010207651,0.0012179345,0.00038507962,0.000632629,0.0004114541,0.00023344459,0.0007705303],"category_scores_gemma":[0.00058621756,0.00022430428,0.00012161655,0.001473155,0.00029964634,0.00019496774,0.00029863615,0.0001913138,0.00010566921],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000809771,0.00033960168,0.4053482,0.0001829784,0.00016447362,0.0028319743,0.0025467747,0.25238296,0.14402369,0.0023885842,0.0024703415,0.18651073],"study_design_scores_gemma":[0.00005816591,0.00005609542,0.6019688,0.000027949862,0.000037016078,0.00024139402,0.0012963633,0.3835897,0.010081925,0.00024430704,0.0023287267,0.00006961455],"about_ca_topic_score_codex":0.79033196,"about_ca_topic_score_gemma":0.8738498,"teacher_disagreement_score":0.20966804,"about_ca_system_score_codex":0.0019472806,"about_ca_system_score_gemma":0.0040513645,"threshold_uncertainty_score":0.42180562},"labels":[],"label_agreement":null},{"id":"W7035953777","doi":"","title":"Appel à contribution - La physiognomonie entre représentations et interprétations. Transpositions esthétiques et transferts internationaux du XIXe au XXIe siècle","year":2012,"lang":"fr","type":"other","venue":"OpenEdition (OpenEdition)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Context (archaeology); Position (finance)","score_opus":0.014317874039414907,"score_gpt":0.29661052068210764,"score_spread":0.2822926466426927,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7035953777","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019287307,0.009444919,0.013280613,0.028397221,0.009750353,0.00013909045,0.0011402584,0.00035144962,0.9182087],"genre_scores_gemma":[0.20820816,0.0057695163,0.0050542415,0.0018683706,0.0020348295,0.00013968203,0.00085034245,0.00052877626,0.7755461],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985935,0.00038074504,0.00005805599,0.00026566803,0.00054629525,0.000155653],"domain_scores_gemma":[0.99843746,0.00046830144,0.00015330329,0.00021738023,0.0005386964,0.00018482986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012918211,0.00069905765,0.0002464172,0.0017872723,0.0042424416,0.004845336,0.0006947698,0.0013802613,0.03440685],"category_scores_gemma":[0.003660622,0.00025300775,0.0003383468,0.0017376364,0.0057880147,0.0034457452,0.0028305599,0.003191251,0.0062847813],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014110471,0.000037629772,0.0014581821,0.0002638379,0.000017850334,0.0005683243,0.053452026,0.00016764688,0.0026197464,0.5789282,0.24917632,0.11316916],"study_design_scores_gemma":[0.000008903163,0.00002134532,0.0029604437,0.00021813625,0.00000590125,0.000331599,0.0053287633,0.00007707379,0.000985434,0.013551108,0.976493,0.000018204832],"about_ca_topic_score_codex":0.038577095,"about_ca_topic_score_gemma":0.051251482,"teacher_disagreement_score":0.038577095,"about_ca_system_score_codex":0.0045893174,"about_ca_system_score_gemma":0.0043738834,"threshold_uncertainty_score":0.11510235},"labels":[],"label_agreement":null},{"id":"W7036006336","doi":"","title":"ANALYSIS AND RECOMMENDATIONS FOR DEVELOPER LEARNING RESOURCES by","year":2012,"lang":"en","type":"article","venue":"eScholarship@McGill (McGill)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Documentation; Traceability; Internal documentation; Technical documentation; Software documentation; Code (set theory); Perspective (graphical); Interview","score_opus":0.022928544417749233,"score_gpt":0.2642125858122072,"score_spread":0.24128404139445797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7036006336","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64425915,0.0033961383,0.27374986,0.006939406,0.000288825,0.0024539917,0.015799697,0.0072302134,0.04588271],"genre_scores_gemma":[0.69190466,0.0010093165,0.27795377,0.0003094676,0.00007695783,0.0008000681,0.011344317,0.00049798744,0.016103502],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9927429,0.0018496659,0.0005310681,0.0016017288,0.002873595,0.00040106033],"domain_scores_gemma":[0.97009236,0.01582899,0.0017928616,0.003032089,0.008434539,0.00081920216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004152693,0.001014107,0.00075026107,0.010820786,0.001147201,0.003733027,0.0014421694,0.0014212446,0.0061656805],"category_scores_gemma":[0.044473216,0.00059648,0.0010843885,0.0071807257,0.00042741463,0.0049283626,0.0014196445,0.0015508079,0.0031780726],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033945678,0.00054698443,0.23780103,0.00077029807,0.00022114083,0.0005971412,0.0076231407,0.00869802,0.0038674516,0.009167489,0.030646844,0.699721],"study_design_scores_gemma":[0.00018329018,0.00063627283,0.21224758,0.0013509152,0.0006186179,0.00089761673,0.022271136,0.5590601,0.017733626,0.027558396,0.15711834,0.0003242314],"about_ca_topic_score_codex":0.02910798,"about_ca_topic_score_gemma":0.051100373,"teacher_disagreement_score":0.02910798,"about_ca_system_score_codex":0.0027373116,"about_ca_system_score_gemma":0.0030148549,"threshold_uncertainty_score":0.057877123},"labels":[],"label_agreement":null},{"id":"W7036412722","doi":"","title":"AURA: A Hybrid Approach to Identify Framework Evolution","year":2010,"lang":"en","type":"article","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Université de Montréal","keywords":"Aura; Context (archaeology); Identity (music); Automatism (medicine)","score_opus":0.03501851785578945,"score_gpt":0.34080767795497585,"score_spread":0.3057891600991864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7036412722","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11137111,0.002266247,0.79495245,0.00036666088,0.0001508799,0.00082531944,0.0045445957,0.08034302,0.0051796804],"genre_scores_gemma":[0.2880545,0.00035557832,0.70164156,0.00022386223,0.000086904634,0.0005469997,0.0042241155,0.001547529,0.0033189931],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99092746,0.0019430383,0.00097050465,0.0024113278,0.003421305,0.00032629434],"domain_scores_gemma":[0.9777753,0.009267882,0.0037402431,0.003512784,0.005167157,0.000536621],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004793208,0.0018550474,0.0017195794,0.018862529,0.0012374572,0.0032973723,0.0026811287,0.0026730557,0.0019208769],"category_scores_gemma":[0.019216528,0.0008942219,0.0019333924,0.0057765697,0.0007516566,0.0042087617,0.0028694458,0.0013426911,0.0018742041],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008088117,0.00036513442,0.05861547,0.0008679935,0.001025632,0.0004501599,0.0011815602,0.0052965167,0.027442968,0.0020072772,0.0124483,0.88949025],"study_design_scores_gemma":[0.00027874848,0.0012369442,0.098689824,0.00030173836,0.0015072781,0.005994826,0.0011263955,0.7410351,0.084428154,0.014844924,0.049745675,0.0008103696],"about_ca_topic_score_codex":0.006351344,"about_ca_topic_score_gemma":0.011289857,"teacher_disagreement_score":0.018862529,"about_ca_system_score_codex":0.0009102316,"about_ca_system_score_gemma":0.0013296264,"threshold_uncertainty_score":0.02534926},"labels":[],"label_agreement":null},{"id":"W7037660893","doi":"","title":"El control del dolor en la Unidad de Cuidados Paliativos - Plan de cuidados estandarizado","year":2019,"lang":"es","type":"dissertation","venue":"Scientia Insularum Revista de Ciencias Naturales en islas","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Palliative care; Unit (ring theory); Pain control; Quality of life (healthcare); Affect (linguistics); Patient care; Plan (archaeology); Ideal (ethics)","score_opus":0.010115670498208117,"score_gpt":0.2771338666173149,"score_spread":0.26701819611910677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7037660893","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33821964,0.29036865,0.023763224,0.0765119,0.0044369134,0.001305955,0.0014570608,0.0009901565,0.26294658],"genre_scores_gemma":[0.74601877,0.107179396,0.043441813,0.002541027,0.0009239486,0.0005529108,0.00073917536,0.00010370894,0.09849923],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995222,0.00018632507,0.000031071064,0.000050536786,0.00012868458,0.000081172184],"domain_scores_gemma":[0.99888545,0.00020386315,0.0002317369,0.000050518316,0.000358158,0.0002702933],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016968239,0.00036403802,0.0003796067,0.0008497762,0.00091073185,0.0023518149,0.00079183176,0.00085647625,0.007266038],"category_scores_gemma":[0.0017903228,0.00014653422,0.00043615806,0.0008491546,0.0007227761,0.00093303865,0.0010502122,0.0014780288,0.0011023643],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005197404,0.0008891068,0.02256204,0.002268847,0.000075335534,0.0006836616,0.007398898,0.0015384419,0.00323157,0.031772472,0.05927873,0.8697812],"study_design_scores_gemma":[0.00027157879,0.0015933838,0.12114375,0.0055747153,0.00025352993,0.0013594936,0.025221575,0.0028238636,0.0040571676,0.014777971,0.82283705,0.00008595054],"about_ca_topic_score_codex":0.010997351,"about_ca_topic_score_gemma":0.02034299,"teacher_disagreement_score":0.010997351,"about_ca_system_score_codex":0.0024911773,"about_ca_system_score_gemma":0.0053091096,"threshold_uncertainty_score":0.02430737},"labels":[],"label_agreement":null},{"id":"W7037862717","doi":"","title":"Expérience spirituelle et expérience de rétablissement en santé mentale : étude descriptive en vue d'une thérapeuthique spirituelle","year":2018,"lang":"fr","type":"other","venue":"Corpus Université Laval (Université Laval)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Prospection; Life style; Belgica","score_opus":0.010751541642756995,"score_gpt":0.22908268516062916,"score_spread":0.21833114351787217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7037862717","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85352707,0.005001862,0.014354648,0.01389939,0.00039670002,0.00029826976,0.00022918588,0.000076860175,0.11221599],"genre_scores_gemma":[0.9824714,0.0018926732,0.0016354966,0.0010556686,0.000031853124,0.00019882988,0.000051962143,0.000027008557,0.012635076],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.995719,0.0026905658,0.00010547421,0.00028369555,0.00080659805,0.0003946156],"domain_scores_gemma":[0.9947126,0.002830334,0.0005647333,0.00031921963,0.001003192,0.0005699281],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050898623,0.0003016062,0.00046647395,0.0011252011,0.005643915,0.005849287,0.0009889096,0.0013038564,0.0050974814],"category_scores_gemma":[0.009418033,0.0002436158,0.00042504503,0.0011549362,0.013374357,0.0037477093,0.0047646146,0.0037403796,0.00042649155],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000025412539,0.00003992048,0.0033622188,0.0002745408,0.000011662926,0.0002676953,0.94821876,0.000064278305,0.00085310295,0.032073982,0.0012027234,0.013605701],"study_design_scores_gemma":[0.000013727825,0.000073592564,0.01593482,0.00048627195,0.000024587265,0.00050628965,0.88960785,0.00020368159,0.000672213,0.0066024843,0.08582241,0.000052065487],"about_ca_topic_score_codex":0.06460101,"about_ca_topic_score_gemma":0.088008784,"teacher_disagreement_score":0.06460101,"about_ca_system_score_codex":0.0100739915,"about_ca_system_score_gemma":0.010316067,"threshold_uncertainty_score":0.12844998},"labels":[],"label_agreement":null},{"id":"W7037963665","doi":"","title":"Global Multidisciplinary Learning in Construction Education: Lessons from Virtual Collaboration of Building Design Teams","year":2012,"lang":"en","type":"article","venue":"DOAJ (DOAJ: Directory of Open Access Journals)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Multidisciplinary approach; Work (physics); Reflection (computer programming); Discipline; Virtual learning environment; Construction industry; Training (meteorology); Virtual reality","score_opus":0.16009930658990595,"score_gpt":0.5421575789109061,"score_spread":0.38205827232100015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7037963665","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7420831,0.007634984,0.026575718,0.029530143,0.00041810333,0.00042800655,0.00005098136,0.000092220725,0.19318685],"genre_scores_gemma":[0.9775922,0.0037620582,0.009319889,0.0009745085,0.00009817982,0.0001564112,0.000024892779,0.00004163198,0.008030279],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9943452,0.0045832223,0.00006750673,0.00018877811,0.00033135258,0.00048396332],"domain_scores_gemma":[0.9919165,0.006258196,0.00026052762,0.0003254863,0.00024262618,0.0009967017],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043512993,0.00052986393,0.0005051377,0.0008810905,0.00511247,0.006198742,0.002126866,0.0022178318,0.002952651],"category_scores_gemma":[0.0069572395,0.0003140904,0.00050961773,0.0011773832,0.008001624,0.0058861,0.008884004,0.0024869072,0.000511358],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020903008,0.0013840786,0.01470288,0.0016710887,0.000074123745,0.0046600886,0.66220933,0.005093681,0.0011234566,0.11715785,0.01253976,0.17917472],"study_design_scores_gemma":[0.0001569514,0.00091208733,0.014270644,0.001484768,0.00004811907,0.002326588,0.6819978,0.0051444056,0.0013746998,0.12073017,0.17147388,0.00007992179],"about_ca_topic_score_codex":0.0029105006,"about_ca_topic_score_gemma":0.006578466,"teacher_disagreement_score":0.006198742,"about_ca_system_score_codex":0.002400615,"about_ca_system_score_gemma":0.002759951,"threshold_uncertainty_score":0.023012161},"labels":[],"label_agreement":null},{"id":"W7038327493","doi":"","title":"On genocide and settler-colonial violence: Australia in comparative perspective","year":2016,"lang":"en","type":"other","venue":"NOVA (University of Newcastle Australia)","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Genocide; Colonialism; Indigenous; Frontier; Causation; Context (archaeology); Politics; Nexus (standard)","score_opus":0.1319351104352924,"score_gpt":0.3624191994640675,"score_spread":0.23048408902877507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7038327493","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55713975,0.11000851,0.00090230803,0.03197909,0.00072087225,0.00007956316,0.00007439755,0.000011416591,0.29908416],"genre_scores_gemma":[0.890252,0.07169635,0.00063949695,0.005015596,0.000330928,0.000084533596,0.00005023274,0.000022473312,0.031908344],"study_design_codex":"qualitative","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9975515,0.0014648121,0.00007540612,0.00013022538,0.00034512067,0.00043302076],"domain_scores_gemma":[0.9975968,0.001412543,0.00029809214,0.00007219268,0.00031582094,0.00030449693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027021861,0.00035743887,0.0006830423,0.0033240546,0.01274634,0.006142393,0.0010259349,0.0031498845,0.005346676],"category_scores_gemma":[0.00472125,0.00031369063,0.0002566761,0.007328111,0.010851979,0.004599364,0.0057750475,0.0034328224,0.00031544894],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026418458,0.00007703449,0.0040143128,0.00088816683,0.000014602854,0.0025409583,0.82221603,0.000159209,0.00057154417,0.14238757,0.0038730411,0.023231167],"study_design_scores_gemma":[0.0000052978094,0.00011908443,0.060640823,0.002613302,0.000033507335,0.0014509778,0.5955469,0.00027975967,0.00020088877,0.01564712,0.32342204,0.00004031539],"about_ca_topic_score_codex":0.23301083,"about_ca_topic_score_gemma":0.41097736,"teacher_disagreement_score":0.23301083,"about_ca_system_score_codex":0.014871066,"about_ca_system_score_gemma":0.009080204,"threshold_uncertainty_score":0.4633091},"labels":[],"label_agreement":null},{"id":"W7038518318","doi":"","title":"Integrating User Feedback to Enhance Software Quality and User Satisfaction in Mobile Application Development","year":2024,"lang":"en","type":"dissertation","venue":"University Library (University of Saskatchewan)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Computer user satisfaction; Android (operating system); User satisfaction; Software quality; User experience design; Source code; Quality (philosophy); User requirements document; Software development","score_opus":0.007609529604559796,"score_gpt":0.22950792453021931,"score_spread":0.22189839492565952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7038518318","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97967005,0.00039752852,0.014773471,0.00066438154,0.000032616772,0.00031685745,0.00010939664,0.00026877443,0.0037670552],"genre_scores_gemma":[0.9868296,0.00017816352,0.011923568,0.00015206021,0.000022517544,0.00017469487,0.00008315097,0.000036478963,0.0005997894],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9781502,0.013762281,0.0013656341,0.0009362631,0.0050820922,0.0007035484],"domain_scores_gemma":[0.85270184,0.09850054,0.016811108,0.003955622,0.02507889,0.0029520649],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017739432,0.0005437497,0.0006465351,0.0023646983,0.0007745174,0.0029796246,0.00050057244,0.00059904036,0.0011959005],"category_scores_gemma":[0.08716874,0.0003451268,0.00071839493,0.0015689767,0.0007397393,0.0020255733,0.0016415243,0.00090858276,0.0004078469],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012349118,0.0014339759,0.50218886,0.0015974883,0.00029918662,0.00027952413,0.027681194,0.0017736291,0.017967014,0.0007540001,0.0021307166,0.4426595],"study_design_scores_gemma":[0.00010020618,0.0054349992,0.92394954,0.0007423883,0.0004809222,0.00043769428,0.0168672,0.024804614,0.016674677,0.0013252259,0.008964004,0.00021840127],"about_ca_topic_score_codex":0.002007691,"about_ca_topic_score_gemma":0.0037291185,"teacher_disagreement_score":0.017739432,"about_ca_system_score_codex":0.0013854271,"about_ca_system_score_gemma":0.0016670942,"threshold_uncertainty_score":0.09381622},"labels":[],"label_agreement":null},{"id":"W7039109163","doi":"","title":"La stéatose hépatique et ses effets sur la régulation du métabolisme du cholestérol &#13;\\nchez le rat","year":2020,"lang":"fr","type":"dissertation","venue":"Papyrus : Institutional Repository (Université de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research","keywords":"Triglyceride; Cholesterol; Weight gain; Lipogenesis; Gene expression; Lipid metabolism; PCSK9; Energy expenditure","score_opus":0.006015900092450333,"score_gpt":0.19132317746639774,"score_spread":0.1853072773739474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7039109163","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9795305,0.0107857995,0.0009982571,0.0012903323,0.00021333458,0.000046730933,0.00068837084,0.00010581676,0.006340904],"genre_scores_gemma":[0.9373264,0.016584687,0.0014147332,0.00061978505,0.0001830724,0.00008919157,0.0005759401,0.00006415435,0.043142013],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99985504,0.000031511132,0.0000071151317,0.000030354146,0.000031634252,0.000044283956],"domain_scores_gemma":[0.9997197,0.00008022024,0.000051111052,0.000041496205,0.000043358028,0.00006405983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00027240202,0.00035312856,0.00036147062,0.00025079577,0.0003251304,0.0010264565,0.0002256931,0.0006532242,0.007112783],"category_scores_gemma":[0.00040854502,0.00040008346,0.00037276544,0.00024787438,0.00070516195,0.0004366708,0.0002311681,0.0009346509,0.0007552217],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.040507674,0.0014098254,0.004507173,0.00033174834,0.00023746376,0.00042702162,0.00029887602,0.00045198086,0.8699159,0.0024422775,0.0016331443,0.07783689],"study_design_scores_gemma":[0.0011856214,0.011988972,0.06434717,0.00008905349,0.00047156718,0.0005014361,0.00052608375,0.0013217698,0.8986382,0.00112684,0.019702425,0.00010083319],"about_ca_topic_score_codex":0.011818455,"about_ca_topic_score_gemma":0.015252137,"teacher_disagreement_score":0.011818455,"about_ca_system_score_codex":0.0006646997,"about_ca_system_score_gemma":0.00067295926,"threshold_uncertainty_score":0.023794591},"labels":[],"label_agreement":null},{"id":"W7039678760","doi":"","title":"Mumbouli Live at The Redfish on 2007-09-21","year":2007,"lang":"en","type":"other","venue":"Bulletin of Miscellaneous Information (Royal Gardens Kew)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Blues; White (mutation); Crypsis; Warbler; Front (military); Nova scotia","score_opus":0.007671899742823482,"score_gpt":0.2045637069346537,"score_spread":0.19689180719183022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7039678760","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026435882,0.0006157154,0.00041038555,0.002836048,0.001975032,0.0001550193,0.006621822,0.0050008763,0.9797416],"genre_scores_gemma":[0.0018495109,0.00013413728,0.00015210724,0.0001317049,0.00004145488,0.000020547623,0.0009243226,0.00027000008,0.9964761],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9998567,0.000011363754,0.0000022519716,0.000021610651,0.00006392439,0.00004411647],"domain_scores_gemma":[0.9996232,0.00002588712,0.00001332598,0.000028338534,0.00012155222,0.00018762228],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00032106135,0.00062332145,0.00040013762,0.0009997516,0.0030118006,0.0020232757,0.0008166587,0.00074782094,0.8343071],"category_scores_gemma":[0.0011026543,0.00030620894,0.00026243742,0.0008688473,0.00032142672,0.0011576384,0.0021120994,0.0009755354,0.6410984],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005326325,0.000012757629,0.00023609844,0.000040841584,9.687797e-7,0.000050370858,0.00009242036,0.0000070821593,0.00020371367,0.00039701728,0.97552,0.023385448],"study_design_scores_gemma":[0.000004590949,0.000013420859,0.000975253,0.000034763885,8.1859645e-7,0.000020349755,0.0001810424,0.000016253836,0.00012222843,0.000070773865,0.99855727,0.0000032532791],"about_ca_topic_score_codex":0.01202244,"about_ca_topic_score_gemma":0.06522335,"teacher_disagreement_score":0.16569293,"about_ca_system_score_codex":0.0011878398,"about_ca_system_score_gemma":0.0008451423,"threshold_uncertainty_score":0.23634076},"labels":[],"label_agreement":null},{"id":"W7045490487","doi":"","title":"Avoin lähdekoodi käytössä: Näkökulmia avoimen lähdekoodin ohjelmistojen käyttöön ja kehittämiseen julkisella ja yksityisellä sektorilla","year":2012,"lang":"fi","type":"article","venue":"Tampere University Institutional Repository (Tampere University)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Work (physics); Order (exchange); Set (abstract data type); Quarter (Canadian coin)","score_opus":0.015621408978008185,"score_gpt":0.19842583107021744,"score_spread":0.18280442209220926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7045490487","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02034897,0.010462098,0.0365961,0.045985978,0.027304424,0.0002856835,0.0021291296,0.010266843,0.8466207],"genre_scores_gemma":[0.05633635,0.006054675,0.020865275,0.0033058885,0.0021317182,0.00012864356,0.002152835,0.0046889735,0.9043355],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.9967801,0.00042146846,0.00015935456,0.000612607,0.0016521013,0.00037433204],"domain_scores_gemma":[0.99281013,0.0010836828,0.000319113,0.0009565915,0.0030148092,0.001815715],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.0027934785,0.0008441625,0.00046366133,0.0013097478,0.0031932401,0.015750173,0.0013311945,0.0019656706,0.13674961],"category_scores_gemma":[0.0084681185,0.0005451028,0.0006214251,0.0016107016,0.0021884108,0.009463579,0.005748949,0.0035295866,0.08743563],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024334046,0.000120517536,0.0022427347,0.000620065,0.00002402747,0.0007476588,0.006043384,0.0002764804,0.010260318,0.050899923,0.65594536,0.27257618],"study_design_scores_gemma":[0.000004802472,0.000012988495,0.00067434215,0.00010601977,0.00000771215,0.00017319247,0.0010425124,0.00021561966,0.0014192528,0.0023707752,0.993956,0.000016829566],"about_ca_topic_score_codex":0.006405445,"about_ca_topic_score_gemma":0.013869608,"teacher_disagreement_score":0.9986688,"about_ca_system_score_codex":0.0030462458,"about_ca_system_score_gemma":0.0046709804,"threshold_uncertainty_score":0.45747304},"labels":[],"label_agreement":null},{"id":"W70716443","doi":"","title":"Proceedings of the 4th international workshop on Predictor models in software engineering","year":2008,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"IBM; Presentation (obstetrics); Software engineering; Computer science; Engineering management; Engineering; Data science","score_opus":0.023690096424751297,"score_gpt":0.23482928284807578,"score_spread":0.21113918642332447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W70716443","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008053617,0.05336087,0.82964754,0.029259276,0.02191151,0.00033277582,0.0016241709,0.0026019714,0.053208195],"genre_scores_gemma":[0.11672721,0.07107471,0.567717,0.0092063565,0.025633438,0.0011655767,0.009440263,0.0030764185,0.19595894],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9948508,0.002524961,0.0003828958,0.0007299312,0.0012268182,0.00028452935],"domain_scores_gemma":[0.9881329,0.0075709606,0.0003111408,0.001422346,0.0019359242,0.00062675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01166815,0.0014037892,0.0018238852,0.0015826209,0.0007693461,0.0052230856,0.0026397917,0.0023596988,0.033790946],"category_scores_gemma":[0.020521495,0.0009427996,0.0025260588,0.0017970392,0.0015859199,0.0054021785,0.003143574,0.0067818207,0.0094682565],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003114273,0.00044963812,0.0021939084,0.000789636,0.00034944474,0.00033463698,0.00063770614,0.01668486,0.0012905232,0.07377443,0.3923224,0.5108614],"study_design_scores_gemma":[0.000081546816,0.00026306318,0.0024373739,0.00087796047,0.00014386054,0.00053818437,0.0003305381,0.06792931,0.0013303998,0.13185762,0.79411745,0.00009263255],"about_ca_topic_score_codex":0.0031975368,"about_ca_topic_score_gemma":0.0034111512,"teacher_disagreement_score":0.033790946,"about_ca_system_score_codex":0.0014715841,"about_ca_system_score_gemma":0.0027562159,"threshold_uncertainty_score":0.113042},"labels":[],"label_agreement":null},{"id":"W7071953711","doi":"","title":"Towards Sustainable AI for Continuous Integration Quality Gates","year":2025,"lang":"en","type":"dissertation","venue":"QSpace (Queen's University Library)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Retraining; Software deployment; Software; Analytics; Quality (philosophy); Process (computing); Scheduling (production processes)","score_opus":0.00862043603205215,"score_gpt":0.24502281688283395,"score_spread":0.2364023808507818,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7071953711","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032242768,0.0011272858,0.9502895,0.0049920543,0.00010002769,0.00026087128,0.00016782507,0.0031036614,0.007716167],"genre_scores_gemma":[0.52841043,0.0013097228,0.46213332,0.0013025859,0.00013502827,0.00053074036,0.00059431535,0.00086995726,0.00471382],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9869811,0.0058137435,0.0006686197,0.0024970677,0.0030920082,0.00094744965],"domain_scores_gemma":[0.95442194,0.030201014,0.0031833118,0.0065547046,0.0042368774,0.0014020832],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01213324,0.0014966064,0.0012462818,0.001562034,0.0010476345,0.006492651,0.00417493,0.0027161788,0.0038293574],"category_scores_gemma":[0.055134624,0.001153487,0.0011666517,0.001747395,0.0039207465,0.010871827,0.009115859,0.005885542,0.0016628996],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021174856,0.00046574714,0.008445532,0.0010628462,0.00020739935,0.00027725362,0.0020268427,0.46568927,0.008710024,0.18326886,0.009176773,0.32045776],"study_design_scores_gemma":[0.000044982226,0.00015249076,0.00062401715,0.000118571836,0.000037262955,0.00009199918,0.00043591906,0.82752746,0.003674174,0.15516609,0.012094034,0.00003290835],"about_ca_topic_score_codex":0.0036510173,"about_ca_topic_score_gemma":0.003594065,"teacher_disagreement_score":0.01213324,"about_ca_system_score_codex":0.002664855,"about_ca_system_score_gemma":0.0049869926,"threshold_uncertainty_score":0.0641675},"labels":[],"label_agreement":null},{"id":"W7084029371","doi":"10.6084/m9.figshare.c.8064939.v1","title":"Knowledge, attitudes, and behaviours towards smoking among people with migration experience: a global scoping review","year":2025,"lang":"en","type":"other","venue":"Figshare","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Focus group; Content analysis; Public health; Qualitative research; Tobacco use; Descriptive statistics; Qualitative property; Data collection","score_opus":0.032296214675812084,"score_gpt":0.3460188292950282,"score_spread":0.31372261461921613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7084029371","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041353023,0.99261945,0.0002557289,0.0010056269,0.00018125058,0.0002706138,0.0003460167,0.0000068896766,0.0011790501],"genre_scores_gemma":[0.015415042,0.9823731,0.0006019578,0.0005498783,0.00010212753,0.00051813765,0.0002681897,0.0000044654626,0.00016703985],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9925965,0.0029006754,0.0024107776,0.0005601106,0.0012128981,0.00031899806],"domain_scores_gemma":[0.9671409,0.025150862,0.0034942674,0.0004191897,0.0034688525,0.00032591046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011725328,0.0010398147,0.0035919691,0.01716081,0.0012685462,0.0038855025,0.0014029118,0.0026812085,0.0033864505],"category_scores_gemma":[0.041077014,0.00090904685,0.0033977746,0.014790569,0.0012696498,0.003320294,0.0023781238,0.0013460809,0.00045351157],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012937536,0.00007393482,0.005137478,0.7109938,0.00255728,0.00044094215,0.0055501773,0.00024009268,0.00041449297,0.0012922499,0.0062837913,0.26688635],"study_design_scores_gemma":[0.000022680137,0.00010926312,0.009983487,0.9438801,0.006319029,0.00056436996,0.0035784561,0.00006407805,0.00013302846,0.00050428265,0.034806065,0.00003505901],"about_ca_topic_score_codex":0.009966383,"about_ca_topic_score_gemma":0.01690621,"teacher_disagreement_score":0.01716081,"about_ca_system_score_codex":0.002586542,"about_ca_system_score_gemma":0.014240059,"threshold_uncertainty_score":0.06201017},"labels":[],"label_agreement":null},{"id":"W7091729","doi":"10.1097/00000542-198207000-00027","title":"Software maintenance maturity model (S3mDSS): a decision support system","year":2008,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec","funders":"","keywords":"Capability Maturity Model; Software maintenance; Software engineering; Computer science; Maturity (psychological); Software development; Decision support system; Software; Task (project management); Personal software process; Software system; Software construction; Process management; Systems engineering; Engineering; Data mining; Operating system","score_opus":0.013080057892800081,"score_gpt":0.22612207024159048,"score_spread":0.2130420123487904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7091729","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06479141,0.0011208153,0.881216,0.006374594,0.00014854565,0.003090233,0.006332453,0.018174155,0.018751781],"genre_scores_gemma":[0.18182062,0.00047582662,0.8109789,0.0002609765,0.000030354011,0.0011227817,0.0041063735,0.00015729346,0.0010469388],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9947848,0.001915857,0.0010859193,0.00040132605,0.0015554298,0.0002567158],"domain_scores_gemma":[0.98195076,0.0081238495,0.0031718437,0.00091792416,0.0049658157,0.00086988165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010842289,0.0012038338,0.0008184766,0.008472287,0.0012392811,0.004544711,0.001883679,0.0021936418,0.0024113206],"category_scores_gemma":[0.03727711,0.000597632,0.0013378558,0.004673735,0.00051586545,0.00718111,0.0027850883,0.002021956,0.0014898847],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006307888,0.001003929,0.056008644,0.0017362324,0.0005566569,0.00048117203,0.0040415837,0.05448095,0.008433014,0.08547083,0.041057043,0.7460991],"study_design_scores_gemma":[0.0005812204,0.0018424568,0.027330402,0.0018644906,0.0009877066,0.0009819198,0.0033479827,0.60346913,0.017131329,0.15112689,0.19077517,0.000561373],"about_ca_topic_score_codex":0.0062641767,"about_ca_topic_score_gemma":0.0059361625,"teacher_disagreement_score":0.010842289,"about_ca_system_score_codex":0.003105403,"about_ca_system_score_gemma":0.0051697143,"threshold_uncertainty_score":0.057340205},"labels":[],"label_agreement":null},{"id":"W7093319411","doi":"10.1016/j.iot.2025.101803","title":"A comparison of code quality metrics and best practices in non-IoT and IoT systems","year":2025,"lang":"en","type":"article","venue":"Internet of Things","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; École de Technologie Supérieure","funders":"","keywords":"Code refactoring; Best practice; Software quality; Software; Software metric; Software system; Metric (unit); Quality (philosophy); Code smell","score_opus":0.08779927866798222,"score_gpt":0.4233576168808828,"score_spread":0.3355583382129006,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7093319411","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94822574,0.0014731735,0.037845343,0.0005362451,0.00006271956,0.0007133958,0.0030377673,0.0012754363,0.006830171],"genre_scores_gemma":[0.86281973,0.0009897838,0.121590674,0.00017454951,0.00002280604,0.0011789225,0.010704958,0.00093024416,0.0015883186],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97236776,0.006877879,0.004041874,0.0021966791,0.013730129,0.00078567397],"domain_scores_gemma":[0.8246571,0.07948296,0.023019247,0.018245416,0.052312948,0.0022824153],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015258213,0.0004612566,0.0005699679,0.014330159,0.000829315,0.0025668032,0.0011932151,0.00062554755,0.001003939],"category_scores_gemma":[0.12329148,0.00043405851,0.0010676769,0.011702012,0.001389483,0.002466571,0.002218008,0.0008846232,0.0003945194],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015632553,0.0012326294,0.3627056,0.0058794073,0.000989265,0.00073442946,0.015593601,0.0183467,0.028158046,0.011098847,0.010483401,0.54321474],"study_design_scores_gemma":[0.00033723132,0.002549026,0.78978217,0.002364015,0.0007958746,0.0012697426,0.0119272405,0.08090808,0.044757124,0.016693186,0.048293237,0.00032307155],"about_ca_topic_score_codex":0.004941041,"about_ca_topic_score_gemma":0.008065806,"teacher_disagreement_score":0.9847418,"about_ca_system_score_codex":0.0027875635,"about_ca_system_score_gemma":0.0033048228,"threshold_uncertainty_score":0.08069408},"labels":[],"label_agreement":null},{"id":"W7093648404","doi":"","title":"Food deserts in Winnipeg, Canada: a novel method for\\nmeasuring a complex and contested construct","year":2017,"lang":"en","type":"article","venue":"PubMed Central","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Construct (python library); Natural (archaeology); Food supply; Food security","score_opus":0.06555332872158774,"score_gpt":0.27304058999486436,"score_spread":0.2074872612732766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7093648404","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6780968,0.004160423,0.10655022,0.0015559023,0.00022575904,0.0015887443,0.07207849,0.0021628907,0.13358091],"genre_scores_gemma":[0.81645864,0.0015480603,0.1451493,0.0001250617,0.000017489212,0.00046754951,0.014769758,0.0003085605,0.021155538],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99969685,0.000031555344,0.000017419072,0.00008085041,0.00011099677,0.00006246707],"domain_scores_gemma":[0.99933213,0.00014471827,0.00008460918,0.000036032434,0.00032156546,0.00008103386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00046144,0.00041045848,0.00034499017,0.00445984,0.002190116,0.0019677184,0.0011076655,0.00036837958,0.0068010027],"category_scores_gemma":[0.00248428,0.00021849065,0.0004002898,0.008567973,0.00060257263,0.0004585926,0.0014204998,0.00029932286,0.00050488935],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006436373,0.00009288025,0.32624248,0.0025977069,0.00045246046,0.0023906012,0.019117454,0.019133482,0.011043029,0.045126755,0.09071779,0.48244184],"study_design_scores_gemma":[0.000085580556,0.00006579122,0.55352336,0.0006509269,0.0003723812,0.00056724204,0.052938275,0.03745458,0.004462589,0.013170452,0.33654568,0.00016315504],"about_ca_topic_score_codex":0.9321012,"about_ca_topic_score_gemma":0.9838748,"teacher_disagreement_score":0.06789881,"about_ca_system_score_codex":0.0055058054,"about_ca_system_score_gemma":0.015088506,"threshold_uncertainty_score":0.13659734},"labels":[],"label_agreement":null},{"id":"W7095139477","doi":"","title":"IMPROVING THE ESTIMATION, CONTINGENCY PLANNING AND TRACKING","year":2010,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Tracking (education); Jury; Emerging technologies; Expert system; Tracking system; Automation","score_opus":0.015534924641010335,"score_gpt":0.27201467262677936,"score_spread":0.256479747985769,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7095139477","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05262912,0.0013230535,0.9316,0.0013943267,0.00022178952,0.00015722105,0.0011193381,0.003356303,0.008198881],"genre_scores_gemma":[0.5729673,0.00085425813,0.41717613,0.00019090802,0.00012626231,0.00013627937,0.0020618353,0.0003219604,0.006165088],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984837,0.00037784138,0.00009581904,0.0004487628,0.0004337707,0.00016008427],"domain_scores_gemma":[0.9953088,0.0025276057,0.00039567825,0.0005967744,0.0009547789,0.00021631151],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018946221,0.0010786441,0.0012139035,0.0026300927,0.0010332828,0.0024439876,0.0013383471,0.001329353,0.0075210137],"category_scores_gemma":[0.012820695,0.0007766803,0.0009242002,0.0023409964,0.000540719,0.0034175974,0.0016724293,0.0014991959,0.0019029686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034883962,0.00019621481,0.016936108,0.00023577898,0.00013021764,0.00021025226,0.0002540694,0.46714875,0.0032220727,0.012265929,0.013703676,0.48534817],"study_design_scores_gemma":[0.000017811637,0.00007265585,0.0027646846,0.00003753978,0.00004166854,0.000081715814,0.00012443679,0.9782259,0.0017044484,0.0128182415,0.004085921,0.000024967043],"about_ca_topic_score_codex":0.03653204,"about_ca_topic_score_gemma":0.033863686,"teacher_disagreement_score":0.03653204,"about_ca_system_score_codex":0.0013137663,"about_ca_system_score_gemma":0.003964388,"threshold_uncertainty_score":0.07263881},"labels":[],"label_agreement":null},{"id":"W7100274705","doi":"","title":"Information Technology de l’information Software Cost Estimation and Control","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Information technology; Identifier; Software; Table (database); Control (management); Information system","score_opus":0.004800158921368816,"score_gpt":0.23929130252562403,"score_spread":0.23449114360425521,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7100274705","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007661577,0.021986198,0.5969968,0.011180577,0.0036167211,0.00074304984,0.006568905,0.007003382,0.3442428],"genre_scores_gemma":[0.23008265,0.04047897,0.35120213,0.002400865,0.002658761,0.0015312243,0.022650978,0.0032885796,0.34570584],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9877494,0.00333959,0.00074087974,0.0012421294,0.0064113704,0.0005166195],"domain_scores_gemma":[0.9828524,0.0042078444,0.00094849244,0.0052286424,0.006401555,0.00036107647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007866968,0.0016587521,0.0013533237,0.0051678782,0.0009649806,0.008396342,0.0014625618,0.0016791917,0.03824241],"category_scores_gemma":[0.031045653,0.0007234087,0.0011484897,0.007209194,0.0016227228,0.006344528,0.002509793,0.0021701078,0.024411686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016002209,0.000054221033,0.0015926096,0.00042831808,0.00007681664,0.000099014236,0.00015939715,0.012617661,0.0013075086,0.1746819,0.093883775,0.71493876],"study_design_scores_gemma":[0.000070676855,0.00013703863,0.00432491,0.00066628604,0.000098890174,0.0002919224,0.00013796153,0.061960857,0.005647418,0.092778966,0.8337863,0.000098765195],"about_ca_topic_score_codex":0.030535892,"about_ca_topic_score_gemma":0.0089119775,"teacher_disagreement_score":0.03824241,"about_ca_system_score_codex":0.0058005354,"about_ca_system_score_gemma":0.008379451,"threshold_uncertainty_score":0.12793356},"labels":[],"label_agreement":null},{"id":"W7100759598","doi":"","title":"in Nunavut: prelude to a screening strategy","year":2015,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.06122732403069267,"score_gpt":0.31845452712967665,"score_spread":0.25722720309898395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7100759598","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.087952316,0.0042508035,0.33654732,0.048621044,0.0022781554,0.0016277034,0.0011857309,0.0009060061,0.5166309],"genre_scores_gemma":[0.6573377,0.0013339121,0.19991808,0.014411352,0.00016183856,0.0015653989,0.0005348994,0.0005931331,0.12414365],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98811615,0.0059127053,0.00046891748,0.0016632228,0.0023428015,0.0014962418],"domain_scores_gemma":[0.9749297,0.014863363,0.0006264994,0.0025567452,0.006053317,0.0009702478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014441967,0.0008899099,0.0014574627,0.0031883155,0.015857032,0.00846512,0.0037683288,0.0037376143,0.0121234],"category_scores_gemma":[0.060515694,0.0008063856,0.0008254974,0.0033246723,0.0068061454,0.008190656,0.0072729373,0.0034944522,0.0020654583],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016992136,0.00009453652,0.008692468,0.0002263509,0.000028142933,0.00045675738,0.009108021,0.00084935804,0.0007730927,0.8433921,0.024741251,0.111468054],"study_design_scores_gemma":[0.000089686,0.0002033469,0.009017865,0.00153636,0.00015675984,0.0008410926,0.029631553,0.014149084,0.007503697,0.64443296,0.29218397,0.00025364608],"about_ca_topic_score_codex":0.2424959,"about_ca_topic_score_gemma":0.5107217,"teacher_disagreement_score":0.7575041,"about_ca_system_score_codex":0.008740788,"about_ca_system_score_gemma":0.029729433,"threshold_uncertainty_score":0.4821688},"labels":[],"label_agreement":null},{"id":"W7103104721","doi":"","title":"Compiler.next: A Search-Based Compiler to Power the AI-Native Future of Software Engineering","year":2025,"lang":"","type":"article","venue":"ArXiv.org","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Software development; Software construction; Compiler; Social software engineering; Interoperability; Software system; Resource-oriented architecture; Software evolution; Cornerstone","score_opus":0.02491462547407967,"score_gpt":0.2845692658986576,"score_spread":0.2596546404245779,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7103104721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0068690707,0.00029276227,0.9669748,0.00073136145,0.00018298122,0.0001294651,0.000118391275,0.017262716,0.007438542],"genre_scores_gemma":[0.07028757,0.00026482972,0.9207989,0.0005612764,0.00007790615,0.00019635052,0.00034426,0.0034811927,0.003987713],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99898547,0.00033900145,0.00008096435,0.0001535786,0.00033355545,0.00010739309],"domain_scores_gemma":[0.997198,0.0013118585,0.0001602347,0.0007485327,0.0004365961,0.00014479064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021091166,0.00071206153,0.00051345315,0.00071459304,0.0007281281,0.0019361168,0.0015978538,0.0010431758,0.0036167384],"category_scores_gemma":[0.007274621,0.0006726146,0.0009549687,0.0004957091,0.0017661275,0.0033621453,0.002719618,0.0020877935,0.0015387996],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046534417,0.00029818044,0.0035007016,0.00069154537,0.00013813184,0.0003719025,0.001313068,0.094434254,0.028061623,0.33716896,0.061335374,0.47222087],"study_design_scores_gemma":[0.00017219695,0.0002366557,0.00045250324,0.0001877717,0.00009199738,0.00035203298,0.00016447296,0.634089,0.022518463,0.20991124,0.1317412,0.000082476305],"about_ca_topic_score_codex":0.0016830813,"about_ca_topic_score_gemma":0.0035465404,"teacher_disagreement_score":0.0036167384,"about_ca_system_score_codex":0.0007556702,"about_ca_system_score_gemma":0.0030762763,"threshold_uncertainty_score":0.012099206},"labels":[],"label_agreement":null},{"id":"W7113415863","doi":"","title":"Large Language Models for Code Generation and Program Comprehension: Exploring Capabilities, Context, and Developer Adaptation","year":2025,"lang":"en","type":"article","venue":"University Library (University of Saskatchewan)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Program comprehension; Documentation; Adaptation (eye); Automatic summarization; Comprehension; Code (set theory); Software development; Java","score_opus":0.031077415142713047,"score_gpt":0.21230285653286524,"score_spread":0.18122544139015218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7113415863","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6101594,0.00078311644,0.37651917,0.0014678722,0.00003800786,0.0004553427,0.00040601712,0.003108311,0.007062753],"genre_scores_gemma":[0.87706304,0.00025152668,0.12022556,0.00015201562,0.000014927019,0.00037393573,0.00061812653,0.00045411361,0.00084682537],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99188775,0.0058494224,0.0002720571,0.0010359463,0.0007918276,0.00016292439],"domain_scores_gemma":[0.9161932,0.071232915,0.0042364234,0.0053344374,0.002332602,0.0006704629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008288639,0.00083708396,0.00044973948,0.0013627373,0.00069799414,0.0039154855,0.0011710888,0.00087113405,0.0023460544],"category_scores_gemma":[0.07133682,0.0007173031,0.00083695276,0.000774054,0.0016279798,0.0076501663,0.0028513924,0.0018706403,0.0004964018],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012711098,0.001434173,0.1795414,0.0016973412,0.00039879399,0.0010952939,0.09742483,0.07291225,0.03402988,0.070900366,0.006452785,0.5328417],"study_design_scores_gemma":[0.00018598234,0.0010118036,0.03990044,0.0006359971,0.0003267513,0.0008113983,0.019302294,0.8188896,0.016828801,0.073440954,0.028408919,0.000257025],"about_ca_topic_score_codex":0.0044579366,"about_ca_topic_score_gemma":0.0050662486,"teacher_disagreement_score":0.008288639,"about_ca_system_score_codex":0.0015537773,"about_ca_system_score_gemma":0.0017658074,"threshold_uncertainty_score":0.043835044},"labels":[],"label_agreement":null},{"id":"W7115741398","doi":"10.2139/ssrn.5932828","title":"Improving IR-based Bug Localization Leveraging Texts and Multimedia from Bug Reports","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; University of Saskatchewan","funders":"","keywords":"Key (lock); Software; Security bug; Application programming interface; Software bug; Interface (matter)","score_opus":0.009396848425966718,"score_gpt":0.248796351078856,"score_spread":0.23939950265288928,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7115741398","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35916263,0.008202006,0.53317195,0.0017086894,0.001534421,0.0008422872,0.008671483,0.072286636,0.014419783],"genre_scores_gemma":[0.6007595,0.0016495491,0.36970976,0.0007092018,0.00137707,0.0003621533,0.012616541,0.0019862682,0.010829961],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968168,0.0006016862,0.0002555009,0.00074127724,0.0013818393,0.00020283769],"domain_scores_gemma":[0.9849576,0.0070265443,0.002161377,0.0015366677,0.0038212202,0.0004965803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018376863,0.0022243,0.0016925499,0.009047736,0.0005213286,0.0017166792,0.0013900878,0.0016480248,0.0033653218],"category_scores_gemma":[0.017812762,0.00047798525,0.00091532466,0.0038527437,0.0003425226,0.003179558,0.0017787718,0.0014153612,0.004764554],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010443922,0.00085988344,0.013890818,0.0012879623,0.000187442,0.0006396074,0.00039988273,0.008108042,0.1080481,0.0005388929,0.021455914,0.84353906],"study_design_scores_gemma":[0.0002970799,0.002930518,0.047719717,0.0003307368,0.0009843148,0.0016646533,0.0010700299,0.7590679,0.14745805,0.0058387374,0.032349687,0.00028865458],"about_ca_topic_score_codex":0.0026740055,"about_ca_topic_score_gemma":0.0037689034,"teacher_disagreement_score":0.009047736,"about_ca_system_score_codex":0.0003226825,"about_ca_system_score_gemma":0.0009411461,"threshold_uncertainty_score":0.011258066},"labels":[],"label_agreement":null},{"id":"W7116086605","doi":"10.5281/zenodo.17970073","title":"OSSVul - ReplicationPackage","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Scripting language; Software; Vulnerability (computing); Artifact (error); Component (thermodynamics); Identification (biology); Data collection; Timestamp","score_opus":0.021157852431734268,"score_gpt":0.25365114832424446,"score_spread":0.2324932958925102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7116086605","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0332162,0.0006891321,0.09142034,0.0007976625,0.00040781117,0.002146033,0.5164535,0.33443996,0.020429319],"genre_scores_gemma":[0.067750834,0.00039006877,0.087038,0.00028928378,0.000083237246,0.0030298876,0.8162466,0.020686788,0.0044853375],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99687237,0.0008042253,0.00042944116,0.0007663553,0.000930067,0.00019755602],"domain_scores_gemma":[0.99005353,0.0026915174,0.00046107845,0.004686253,0.0018350896,0.00027251808],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0053053233,0.0020758347,0.0010558782,0.0035473695,0.0007576448,0.0028145062,0.0031940069,0.00087174523,0.020869128],"category_scores_gemma":[0.022722533,0.0009504547,0.0020562068,0.00300499,0.000642816,0.003092003,0.0029719637,0.0017127906,0.014139712],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011438663,0.0005673841,0.019199828,0.002495756,0.0004707132,0.0002501953,0.0006509668,0.052240428,0.0037991055,0.01627783,0.8017792,0.10112464],"study_design_scores_gemma":[0.0008812541,0.0006952852,0.019056078,0.00063588924,0.00018928818,0.00037035844,0.0007088819,0.24418452,0.0138465185,0.0324896,0.6866734,0.00026890164],"about_ca_topic_score_codex":0.0118781235,"about_ca_topic_score_gemma":0.009895256,"teacher_disagreement_score":0.99469465,"about_ca_system_score_codex":0.0012845901,"about_ca_system_score_gemma":0.0022471072,"threshold_uncertainty_score":0.069814205},"labels":[],"label_agreement":null},{"id":"W7116107539","doi":"10.5281/zenodo.17970072","title":"OSSVul - ReplicationPackage","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Scripting language; Software; Vulnerability (computing); Artifact (error); Component (thermodynamics); Identification (biology); Data collection; Timestamp","score_opus":0.021157852431734268,"score_gpt":0.25365114832424446,"score_spread":0.2324932958925102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7116107539","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0332162,0.0006891321,0.09142034,0.0007976625,0.00040781117,0.002146033,0.5164535,0.33443996,0.020429319],"genre_scores_gemma":[0.067750834,0.00039006877,0.087038,0.00028928378,0.000083237246,0.0030298876,0.8162466,0.020686788,0.0044853375],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99687237,0.0008042253,0.00042944116,0.0007663553,0.000930067,0.00019755602],"domain_scores_gemma":[0.99005353,0.0026915174,0.00046107845,0.004686253,0.0018350896,0.00027251808],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053053233,0.0020758347,0.0010558782,0.0035473695,0.0007576448,0.0028145062,0.0031940069,0.00087174523,0.020869128],"category_scores_gemma":[0.022722533,0.0009504547,0.0020562068,0.00300499,0.000642816,0.003092003,0.0029719637,0.0017127906,0.014139712],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011438663,0.0005673841,0.019199828,0.002495756,0.0004707132,0.0002501953,0.0006509668,0.052240428,0.0037991055,0.01627783,0.8017792,0.10112464],"study_design_scores_gemma":[0.0008812541,0.0006952852,0.019056078,0.00063588924,0.00018928818,0.00037035844,0.0007088819,0.24418452,0.0138465185,0.0324896,0.6866734,0.00026890164],"about_ca_topic_score_codex":0.0118781235,"about_ca_topic_score_gemma":0.009895256,"teacher_disagreement_score":0.020869128,"about_ca_system_score_codex":0.0012845901,"about_ca_system_score_gemma":0.0022471072,"threshold_uncertainty_score":0.069814205},"labels":[],"label_agreement":null},{"id":"W7117107227","doi":"10.1145/3756681.3756954","title":"CCCI: Code Completion with Contextual Information for Complex Data Transfer Tasks Using Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Consortium de Recherche et d’innovation en Aérospatiale au Québec","keywords":"Scripting language; Code (set theory); Source lines of code; Java; Table (database); Object (grammar); Source code; Task (project management)","score_opus":0.10474425969865973,"score_gpt":0.3437311531168272,"score_spread":0.23898689341816748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117107227","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016462728,0.00032385107,0.83462584,0.00026080225,0.00011554172,0.00048711195,0.0015189201,0.1446753,0.0015299662],"genre_scores_gemma":[0.11658859,0.00018585804,0.8613678,0.00025687425,0.000064686574,0.00079557917,0.0084964745,0.009919715,0.0023243767],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961922,0.00096805993,0.0002777069,0.00097193505,0.0013650695,0.00022504838],"domain_scores_gemma":[0.98703116,0.0064811064,0.0009873741,0.003022656,0.0020392092,0.00043838643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038740751,0.0026446006,0.0010137748,0.0019190155,0.00084009406,0.0018607696,0.003367997,0.0015192869,0.0063574617],"category_scores_gemma":[0.026692878,0.001439901,0.0020823414,0.0012563762,0.0012355262,0.0033931227,0.0040222663,0.0042126705,0.0042315675],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009828532,0.0008878289,0.011620443,0.0014961475,0.0002848183,0.00065897166,0.002628935,0.13416484,0.032959215,0.011402282,0.08309585,0.71981776],"study_design_scores_gemma":[0.00010556687,0.00016889481,0.0011213206,0.00006023612,0.000052716347,0.00013292737,0.0001967205,0.9570055,0.017934963,0.00718519,0.015962537,0.00007345357],"about_ca_topic_score_codex":0.010973541,"about_ca_topic_score_gemma":0.017742796,"teacher_disagreement_score":0.010973541,"about_ca_system_score_codex":0.0014419291,"about_ca_system_score_gemma":0.0046169776,"threshold_uncertainty_score":0.021819353},"labels":[],"label_agreement":null},{"id":"W7117107742","doi":"10.2139/ssrn.5963911","title":"A First Look at the Self-Admitted Technical Debt in Test Code: Taxonomy and Detection","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Technical debt; Test (biology); Code (set theory); Java; Taxonomy (biology); Software","score_opus":0.01097806144053919,"score_gpt":0.2440139195543555,"score_spread":0.2330358581138163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117107742","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82419825,0.01135442,0.11506737,0.0077356063,0.0001852936,0.00031398417,0.0034377354,0.0013679506,0.03633943],"genre_scores_gemma":[0.9343562,0.0031714435,0.054637376,0.00067084056,0.00012812964,0.00010063161,0.0020863866,0.00017713875,0.004671905],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9955195,0.0007590182,0.00070298335,0.00058041775,0.0019854757,0.00045259725],"domain_scores_gemma":[0.9600352,0.01782848,0.007440087,0.0036305692,0.008975707,0.0020899437],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002706991,0.00038514202,0.0005416607,0.011691461,0.0026106252,0.0040733516,0.0014613611,0.0021988626,0.0022344857],"category_scores_gemma":[0.023467902,0.0003578806,0.0007441626,0.010504806,0.003362842,0.007617745,0.003272108,0.0026878307,0.0007467043],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003043519,0.0003096715,0.61219555,0.0007614896,0.00007226659,0.001461451,0.014375616,0.0016109609,0.012333654,0.113795325,0.0072760424,0.23550373],"study_design_scores_gemma":[0.000049641516,0.0009781023,0.5557202,0.0021812243,0.00023396903,0.017338296,0.02833524,0.048709475,0.019865453,0.20881033,0.117393,0.00038494397],"about_ca_topic_score_codex":0.0051900577,"about_ca_topic_score_gemma":0.0065093897,"teacher_disagreement_score":0.011691461,"about_ca_system_score_codex":0.001454038,"about_ca_system_score_gemma":0.0025308526,"threshold_uncertainty_score":0.014316082},"labels":[],"label_agreement":null},{"id":"W7117109273","doi":"10.1145/3756681.3756955","title":"Unveiling Ruby: Insights from Stack Overflow and Developer Survey","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; Okanagan College","funders":"","keywords":"Key (lock); Quality (philosophy); Stack (abstract data type); Face (sociological concept); Security bug","score_opus":0.027764503320739698,"score_gpt":0.2772876625916046,"score_spread":0.24952315927086488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117109273","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.992871,0.0005875632,0.0025944773,0.0011247679,0.000019419222,0.00007938732,0.00079179666,0.00010012575,0.0018315505],"genre_scores_gemma":[0.99156874,0.00071073446,0.00372223,0.00076338864,0.000045671237,0.00022571083,0.0014504873,0.00010365103,0.0014094637],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9879996,0.005053511,0.001069795,0.0009517509,0.003947905,0.000977422],"domain_scores_gemma":[0.9145904,0.04653294,0.016960068,0.0021283352,0.01656575,0.0032224709],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016082866,0.00041613335,0.00044530988,0.0056394534,0.00093439576,0.001479313,0.0006667347,0.0007205358,0.0008122981],"category_scores_gemma":[0.06964084,0.0003419668,0.00028216856,0.0029267778,0.00067639985,0.00335186,0.0024435713,0.0010661513,0.0003743407],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001330602,0.00021251544,0.68196774,0.00091535354,0.00006823748,0.00076978275,0.17060106,0.0004253616,0.0040456257,0.0010567574,0.012974807,0.1268297],"study_design_scores_gemma":[0.000018081933,0.00032599297,0.79484993,0.00078216125,0.000051685864,0.0008951237,0.15382504,0.0036388848,0.0020317293,0.00080890476,0.04262879,0.00014364184],"about_ca_topic_score_codex":0.0068122493,"about_ca_topic_score_gemma":0.011487809,"teacher_disagreement_score":0.016082866,"about_ca_system_score_codex":0.0016997965,"about_ca_system_score_gemma":0.0016641191,"threshold_uncertainty_score":0.08505535},"labels":[],"label_agreement":null},{"id":"W7117117068","doi":"10.1145/3756681.3756994","title":"BugsRepo: A Comprehensive Curated Dataset of Bug Reports, Comments and Contributors Information from Bugzilla","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Security bug; Software bug; Software; Key (lock); Software maintenance; Software regression","score_opus":0.012944856714436022,"score_gpt":0.2832649274228555,"score_spread":0.27032007070841946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117117068","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03572261,0.0014972031,0.0053103804,0.0005473654,0.00020730284,0.0003493472,0.9372361,0.015378188,0.0037515047],"genre_scores_gemma":[0.01657134,0.00024296787,0.009680734,0.00012796001,0.000041045976,0.00045827343,0.97102904,0.0006234283,0.0012252535],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961267,0.0005969855,0.0006090712,0.0009881465,0.0013795604,0.0002995583],"domain_scores_gemma":[0.9850739,0.003817954,0.0031867346,0.0032697408,0.0032748182,0.001376857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002295996,0.0018797115,0.00085652765,0.010551694,0.0012537118,0.0013937983,0.002330466,0.002407505,0.0035415576],"category_scores_gemma":[0.01735503,0.0007874089,0.0011769914,0.0064931065,0.0007118308,0.001733188,0.003362382,0.0019627644,0.006215985],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082539796,0.00057508313,0.060292,0.004104775,0.00038308502,0.00080693146,0.0013907559,0.0029195887,0.010544458,0.0023111645,0.8397614,0.07608549],"study_design_scores_gemma":[0.00058746355,0.00039356778,0.18597843,0.00057491055,0.00026554696,0.0013264907,0.0005869383,0.010836764,0.008534944,0.0026170257,0.7880029,0.00029499183],"about_ca_topic_score_codex":0.018009104,"about_ca_topic_score_gemma":0.040978447,"teacher_disagreement_score":0.018009104,"about_ca_system_score_codex":0.0012810396,"about_ca_system_score_gemma":0.003393792,"threshold_uncertainty_score":0.035808563},"labels":[],"label_agreement":null},{"id":"W7117118041","doi":"10.1145/3756681.3756971","title":"Reinforcement Learning vs Supervised Learning: A tug of war to generate refactored code accurately","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Code refactoring; Reinforcement learning; Code (set theory); Source code; Java; Software; Suite","score_opus":0.045803679298979987,"score_gpt":0.3140277163149539,"score_spread":0.2682240370159739,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117118041","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060706634,0.0007667468,0.9304304,0.0006804191,0.0001000875,0.00011542555,0.000072982395,0.0049514524,0.002175776],"genre_scores_gemma":[0.7626691,0.000326976,0.2327867,0.00069804426,0.00009237883,0.00018908901,0.00033099396,0.00061742583,0.0022893383],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984005,0.0006573343,0.000075505355,0.00044421866,0.00029338375,0.00012906405],"domain_scores_gemma":[0.99491113,0.0030166574,0.00038177587,0.0010341131,0.0004909901,0.00016536865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027336448,0.0013280325,0.00095097645,0.0006694714,0.00035800907,0.0008220275,0.0015029316,0.0012105117,0.0015612429],"category_scores_gemma":[0.010448014,0.0004845903,0.0006385725,0.00046654802,0.0011264554,0.0016366582,0.0012308969,0.0024652034,0.00073169585],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034862966,0.0002457067,0.0032064172,0.00014177784,0.00008416755,0.00008846394,0.0001535432,0.71357054,0.007755642,0.0051490916,0.0026537487,0.26660234],"study_design_scores_gemma":[0.00001629754,0.000065372886,0.00014328294,0.000009170427,0.0000081218295,0.000015441874,0.00000987096,0.9942742,0.0017278461,0.0032259503,0.0004982204,0.0000061963406],"about_ca_topic_score_codex":0.0040134294,"about_ca_topic_score_gemma":0.0032442468,"teacher_disagreement_score":0.0040134294,"about_ca_system_score_codex":0.00093250175,"about_ca_system_score_gemma":0.0016923969,"threshold_uncertainty_score":0.014457107},"labels":[],"label_agreement":null},{"id":"W7117120591","doi":"10.1145/3756681.3756978","title":"Understanding the Impact of Domain Term Explanation on Duplicate Bug Report Detection","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Jargon; Security bug; Focus (optics); Domain (mathematical analysis); Term (time); Replicate; Software bug","score_opus":0.06766818080844851,"score_gpt":0.33610771122697475,"score_spread":0.26843953041852625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117120591","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90824646,0.0052178036,0.07689861,0.0019195236,0.00013063302,0.00028324517,0.00087359064,0.0034909267,0.002939117],"genre_scores_gemma":[0.92537403,0.00078985485,0.07127044,0.00024939052,0.00006722877,0.00007711022,0.0014879928,0.00017330253,0.0005107723],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9851054,0.007065759,0.0012588517,0.0026469093,0.0031486426,0.0007744175],"domain_scores_gemma":[0.73355937,0.22195134,0.015864473,0.012669896,0.014231099,0.0017237434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015589046,0.0011833739,0.0011389435,0.0055195247,0.00096170936,0.0031227444,0.0013437778,0.0016282427,0.0008146581],"category_scores_gemma":[0.15468541,0.0006742921,0.0010913637,0.0029718038,0.0010477384,0.00831415,0.0026626394,0.002364858,0.0004288461],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012092931,0.0009537714,0.48602957,0.0021899953,0.0005881784,0.00095782266,0.006484872,0.017127695,0.020691061,0.001761918,0.0059320843,0.45607376],"study_design_scores_gemma":[0.000294094,0.0016123062,0.4147658,0.0007012107,0.0017699887,0.0028225826,0.0065876422,0.50655836,0.035921294,0.011370214,0.017223442,0.0003730703],"about_ca_topic_score_codex":0.0091666635,"about_ca_topic_score_gemma":0.01162618,"teacher_disagreement_score":0.015589046,"about_ca_system_score_codex":0.0010725356,"about_ca_system_score_gemma":0.002411121,"threshold_uncertainty_score":0.082443714},"labels":[],"label_agreement":null},{"id":"W7117144342","doi":"10.1145/3756681.3756974","title":"Racing Against the Clock: Exploring the Impact of Scheduled Deadlines on Technical Debt","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Commit; Technical debt; Debt; Software; Coding (social sciences); Software development; Empirical research","score_opus":0.0454042560801059,"score_gpt":0.33571498017053486,"score_spread":0.29031072409042896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117144342","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9914493,0.0002655422,0.0037091938,0.00043166216,0.000019305762,0.00006086364,0.00016664548,0.000031294294,0.0038661913],"genre_scores_gemma":[0.99761945,0.0000956932,0.001795868,0.000052283496,0.0000073587576,0.000038359998,0.00008684685,0.00001888481,0.00028520092],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99179333,0.0036721854,0.00058681326,0.00064814807,0.0026767205,0.000622754],"domain_scores_gemma":[0.80360997,0.122864164,0.04916172,0.0038639747,0.01642511,0.004075115],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012193495,0.00032890486,0.00024564992,0.0032063243,0.0010259027,0.0026217976,0.00073471357,0.00044798624,0.0016726163],"category_scores_gemma":[0.10751141,0.00032560495,0.00031484233,0.002493885,0.0015346315,0.003332047,0.002368695,0.0012695796,0.00021498428],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035260685,0.0002706435,0.85640913,0.0006202203,0.00010568947,0.00081294007,0.046161972,0.0028535058,0.002768737,0.00502087,0.0014899517,0.08313377],"study_design_scores_gemma":[0.00001888954,0.0005796205,0.91961265,0.00040547334,0.000079496625,0.00037086863,0.05739677,0.008157594,0.0014024148,0.0038157231,0.008087895,0.00007263295],"about_ca_topic_score_codex":0.006598321,"about_ca_topic_score_gemma":0.009623644,"teacher_disagreement_score":0.012193495,"about_ca_system_score_codex":0.002587855,"about_ca_system_score_gemma":0.003562658,"threshold_uncertainty_score":0.064486146},"labels":[],"label_agreement":null},{"id":"W7117149526","doi":"10.1145/3756681.3756997","title":"Large Language Models for API Classification: An Explorative Study","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; Université du Québec à Rimouski","funders":"","keywords":"Task (project management); Function (biology); Context (archaeology); Reliability (semiconductor); Software; Resilience (materials science)","score_opus":0.08677487955050701,"score_gpt":0.3737799482582607,"score_spread":0.2870050687077537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117149526","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7997595,0.0071610636,0.16439524,0.0032786052,0.00034427317,0.0012907524,0.0105020115,0.0052463007,0.008022294],"genre_scores_gemma":[0.85247386,0.0010660135,0.122765936,0.0006443956,0.00015726114,0.0012757136,0.019121625,0.0008401603,0.001655053],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97865915,0.015846455,0.0011046724,0.0020848901,0.0017854568,0.00051942316],"domain_scores_gemma":[0.80812305,0.17776239,0.003699906,0.0060599856,0.003565665,0.0007890712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017605757,0.002205586,0.0010582116,0.0036816173,0.0012122298,0.004274899,0.0022702678,0.0015568673,0.0026137498],"category_scores_gemma":[0.071105376,0.00078086066,0.0031859435,0.002927184,0.0011960287,0.0071800933,0.0031116784,0.00492459,0.001844732],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030149727,0.007612511,0.24959745,0.005666952,0.002180941,0.0028209279,0.017887833,0.07552893,0.011058218,0.013975754,0.055996247,0.5546593],"study_design_scores_gemma":[0.00041635556,0.0012353336,0.057425365,0.0011847175,0.0010106807,0.0018721105,0.010648496,0.8552017,0.008747089,0.023689935,0.03825413,0.00031412824],"about_ca_topic_score_codex":0.008428107,"about_ca_topic_score_gemma":0.010890549,"teacher_disagreement_score":0.017605757,"about_ca_system_score_codex":0.002104456,"about_ca_system_score_gemma":0.0018644566,"threshold_uncertainty_score":0.09310925},"labels":[],"label_agreement":null},{"id":"W7117152675","doi":"10.1145/3756681.3756995","title":"Can We Enhance Bug Report Quality Using LLMs?: An Empirical Study of LLM-Based Bug Report Generation","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Triage; Quality (philosophy); Software regression; Empirical research; Software; Software bug","score_opus":0.1264404597549885,"score_gpt":0.45160629580921524,"score_spread":0.3251658360542268,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117152675","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9081673,0.004489121,0.04679954,0.0011395711,0.00016362089,0.00034034124,0.0036503258,0.03343617,0.0018140405],"genre_scores_gemma":[0.93092614,0.00048875844,0.05847355,0.00030486853,0.000042073196,0.00015838302,0.0079507455,0.0008312848,0.0008242306],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9916937,0.00383495,0.0006947433,0.0017873019,0.001703677,0.00028561466],"domain_scores_gemma":[0.93994147,0.03915252,0.0054169153,0.009911431,0.004349294,0.0012283616],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011052822,0.0012070633,0.00071266736,0.0030270147,0.00037767662,0.0015565768,0.0019874906,0.0012847694,0.0007659102],"category_scores_gemma":[0.09064268,0.0005062531,0.00082020793,0.0018632396,0.00090258336,0.0034386872,0.0016218757,0.0019361933,0.00090151676],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014995807,0.0014849048,0.22546834,0.0021661262,0.00049038394,0.00052211253,0.0028643003,0.06604167,0.018028669,0.0010754564,0.025578238,0.65478015],"study_design_scores_gemma":[0.00050411065,0.0031551928,0.11310636,0.00037121974,0.00048702277,0.0013336568,0.001468145,0.82479084,0.030929092,0.0038922324,0.019704452,0.00025756343],"about_ca_topic_score_codex":0.005792437,"about_ca_topic_score_gemma":0.0072583817,"teacher_disagreement_score":0.98894715,"about_ca_system_score_codex":0.0011921404,"about_ca_system_score_gemma":0.0014516162,"threshold_uncertainty_score":0.05845356},"labels":[],"label_agreement":null},{"id":"W7117157423","doi":"10.1145/3756681.3756977","title":"Do Automatic Comment Generation Techniques Fall Short? Exploring the Influence of Method Dependencies on Code Understanding","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Program comprehension; Code (set theory); Software; Java; Source code; Code review; Quality (philosophy); Static program analysis","score_opus":0.16621178447451343,"score_gpt":0.3673553603012008,"score_spread":0.20114357582668738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117157423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36655426,0.0074987034,0.5215304,0.004408956,0.00090141705,0.0009493888,0.0052892715,0.08547018,0.0073974654],"genre_scores_gemma":[0.6614285,0.0014020798,0.31312415,0.0015273129,0.0003871411,0.00068222464,0.010223722,0.005614134,0.005610778],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98429626,0.007975232,0.0009087061,0.0025169877,0.0038113727,0.000491449],"domain_scores_gemma":[0.84426725,0.10965425,0.010837504,0.016584653,0.016844153,0.0018121954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01307756,0.0020168743,0.0011579617,0.0027681484,0.000812727,0.0023531164,0.0021113849,0.0019021553,0.0024803258],"category_scores_gemma":[0.109875135,0.00084855163,0.0011280763,0.0013739556,0.0008474439,0.0065499237,0.0023165827,0.0027511076,0.0033607795],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011979451,0.00059383566,0.060494933,0.0031177334,0.00032678084,0.00053527055,0.0066801994,0.0069591985,0.045760524,0.0018294306,0.042653438,0.8298506],"study_design_scores_gemma":[0.0006602806,0.0026587935,0.09390164,0.00207218,0.00088998576,0.002332807,0.0068399454,0.61596483,0.11134732,0.023438329,0.13922553,0.0006684267],"about_ca_topic_score_codex":0.0025768764,"about_ca_topic_score_gemma":0.005254306,"teacher_disagreement_score":0.01307756,"about_ca_system_score_codex":0.00087977725,"about_ca_system_score_gemma":0.0022481144,"threshold_uncertainty_score":0.069161534},"labels":[],"label_agreement":null},{"id":"W7117166478","doi":"10.1145/3756681.3757017","title":"A Study on Mixup-Inspired Augmentation Methods for Software Vulnerability Detection","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Software; Metric (unit); Source code; Augment; Code (set theory); Vulnerability (computing); Embedding; Software bug","score_opus":0.060454118786172316,"score_gpt":0.4304843882277117,"score_spread":0.3700302694415394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117166478","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16529368,0.0063968007,0.8131659,0.0016852987,0.0003995818,0.00022359958,0.0005003454,0.0068592676,0.0054754536],"genre_scores_gemma":[0.73367,0.0013080544,0.25524122,0.0011592705,0.0002753048,0.00030660117,0.001686929,0.0004943153,0.0058583254],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998697,0.0004555034,0.00007199766,0.00037888094,0.0002860863,0.000110496956],"domain_scores_gemma":[0.99543273,0.002755938,0.00028561332,0.00084081024,0.0005239059,0.00016095424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025224185,0.0018467472,0.0012033082,0.0014900898,0.0004667155,0.0011967613,0.0020633796,0.0014642968,0.0018219816],"category_scores_gemma":[0.009611703,0.00060632353,0.001258031,0.0011584188,0.000981782,0.0037275723,0.0022160755,0.0025577254,0.00072213786],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000760666,0.00060476933,0.011687618,0.00041782716,0.00031110347,0.00028831285,0.0003020591,0.27293134,0.01780136,0.011812031,0.009583626,0.67349935],"study_design_scores_gemma":[0.000012912913,0.00014845526,0.00051531184,0.000027527169,0.000029860634,0.00009403439,0.000022011007,0.98908573,0.0040311725,0.0038996174,0.0021210094,0.0000123479385],"about_ca_topic_score_codex":0.0018770053,"about_ca_topic_score_gemma":0.002107644,"teacher_disagreement_score":0.0025224185,"about_ca_system_score_codex":0.0006820084,"about_ca_system_score_gemma":0.0006953004,"threshold_uncertainty_score":0.013339996},"labels":[],"label_agreement":null},{"id":"W7117252587","doi":"10.1109/vl-hcc65237.2025.00051","title":"Cracking CodeWhisperer: Analyzing Developers’ Interactions and Patterns During Programming Tasks","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Structuring; Software; Code (set theory); Baseline (sea); Abstraction; Process (computing); Natural language; Task analysis","score_opus":0.0188104640193474,"score_gpt":0.2996216054758287,"score_spread":0.28081114145648134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117252587","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99261534,0.00005950164,0.006184602,0.00009438759,0.000003846291,0.000079678124,0.00012762872,0.00020686002,0.00062810513],"genre_scores_gemma":[0.97694117,0.00006109647,0.02099418,0.00007195471,0.000004379076,0.00018953961,0.00044583864,0.000103562816,0.0011883035],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99645007,0.0018457598,0.00019591823,0.00058142364,0.0007271994,0.00019970938],"domain_scores_gemma":[0.95476466,0.033998307,0.0040897964,0.0025177984,0.0037506039,0.00087891874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004278054,0.0004348512,0.0003088073,0.0016529709,0.00067751156,0.0009877039,0.00065676955,0.00072680507,0.0007073791],"category_scores_gemma":[0.035960305,0.00036733068,0.00019371048,0.000925179,0.0007145563,0.0015349566,0.0015024772,0.0007022738,0.00030843567],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072742987,0.0005930454,0.4749618,0.00066632376,0.000097465105,0.0009720352,0.22271329,0.0009658483,0.055199847,0.0008623306,0.0029096692,0.23933086],"study_design_scores_gemma":[0.00010965739,0.0012489925,0.8387688,0.00025559842,0.00010590336,0.0013136055,0.085203156,0.03048003,0.02364277,0.0023545427,0.01627685,0.00024014687],"about_ca_topic_score_codex":0.004223119,"about_ca_topic_score_gemma":0.012125931,"teacher_disagreement_score":0.004278054,"about_ca_system_score_codex":0.00058505544,"about_ca_system_score_gemma":0.00084784883,"threshold_uncertainty_score":0.022624731},"labels":[],"label_agreement":null},{"id":"W7117313584","doi":"10.1007/s10115-025-02657-2","title":"Enhanced software defect prediction using edge feature and self-attention GAN with pelican optimization","year":2025,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Feature (linguistics); Convolutional neural network; Software; Enhanced Data Rates for GSM Evolution; Artificial neural network; Software bug; Pattern recognition (psychology); Software system; Deep learning","score_opus":0.006036464635271571,"score_gpt":0.22966269469882605,"score_spread":0.22362623006355448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117313584","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06730061,0.00052305387,0.9268516,0.00021073956,0.000071748545,0.000034635457,0.00021144778,0.0019008863,0.0028953215],"genre_scores_gemma":[0.8349361,0.00021697501,0.15669139,0.00032483647,0.000072138144,0.00008263584,0.00090149086,0.0002559264,0.0065184105],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997658,0.000038823444,0.000009168431,0.00007573621,0.000075538774,0.00003494032],"domain_scores_gemma":[0.99952364,0.00021869525,0.00003759529,0.000055568173,0.00014138043,0.000023014198],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004572316,0.00078506046,0.0009948313,0.0007615689,0.00018366579,0.00046096413,0.001097354,0.0007684284,0.0016513655],"category_scores_gemma":[0.0010756797,0.00027206147,0.00069003925,0.00054112554,0.00024940036,0.0008731938,0.0004940538,0.00072980725,0.00040648898],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002726526,0.00023217621,0.0034309935,0.00009502385,0.00010625576,0.00012017525,0.000039438517,0.560045,0.019679483,0.0037568954,0.007045634,0.40517634],"study_design_scores_gemma":[0.0000022982451,0.000011431807,0.00019221428,0.0000012602161,0.000005204764,0.000008841039,0.0000013535279,0.9986823,0.0006320789,0.00036154623,0.00009973648,0.0000017745562],"about_ca_topic_score_codex":0.004411426,"about_ca_topic_score_gemma":0.008136703,"teacher_disagreement_score":0.004411426,"about_ca_system_score_codex":0.00039236684,"about_ca_system_score_gemma":0.0005976309,"threshold_uncertainty_score":0.008771479},"labels":[],"label_agreement":null},{"id":"W7117477189","doi":"10.1145/3786771","title":"Assessing and Advancing Benchmarks for Evaluating Large Language Models in Software Engineering Tasks","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Software Engineering Research","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Popularity; Software quality assurance; Model-driven architecture; Software development; Coding (social sciences); Quality (philosophy); Software; Benchmark (surveying); Social software engineering","score_opus":0.07511031672640531,"score_gpt":0.3920066529526886,"score_spread":0.3168963362262833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117477189","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35236132,0.017329024,0.5757956,0.0029476874,0.0010179993,0.0021646393,0.0054942477,0.011877487,0.031011993],"genre_scores_gemma":[0.50213325,0.003706838,0.4746594,0.00044481418,0.00015906994,0.0019688192,0.012947919,0.0018905946,0.0020893642],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.96129864,0.021006433,0.004160275,0.0017950061,0.010619751,0.001119984],"domain_scores_gemma":[0.84954315,0.10046197,0.008795676,0.014303503,0.024599086,0.0022966324],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028474636,0.0021027406,0.0011808324,0.0070931436,0.0010978103,0.0044698333,0.0033217552,0.0015706029,0.0021096505],"category_scores_gemma":[0.15091833,0.00064067804,0.0010085861,0.0071850573,0.001262063,0.00547589,0.0036110946,0.0024662684,0.000944953],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015155736,0.0021974293,0.04428716,0.0054569943,0.00057861686,0.0003161597,0.0022337695,0.21614581,0.015310949,0.058141638,0.03178678,0.6220291],"study_design_scores_gemma":[0.00036959958,0.002782533,0.02373221,0.0027037559,0.00033167773,0.00042917213,0.002290535,0.8092759,0.040398087,0.06564186,0.05172943,0.00031513412],"about_ca_topic_score_codex":0.0067143594,"about_ca_topic_score_gemma":0.009007523,"teacher_disagreement_score":0.9715254,"about_ca_system_score_codex":0.0031144547,"about_ca_system_score_gemma":0.004158453,"threshold_uncertainty_score":0.15059012},"labels":[],"label_agreement":null},{"id":"W7117557885","doi":"10.2139/ssrn.5986423","title":"RAGdeterm: Deterministic Retrieval-Augmented Generation for Code Generator","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Code generation; Code (set theory); Generator (circuit theory); Inheritance (genetic algorithm); Software; Relational database; Source code; Dead code","score_opus":0.03137808466897117,"score_gpt":0.3101581487205192,"score_spread":0.27878006405154804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117557885","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008047038,0.00023535064,0.9356338,0.00020592216,0.00012451003,0.0001194616,0.0005558573,0.051981088,0.0030969423],"genre_scores_gemma":[0.28581488,0.00019108529,0.69378597,0.00045514884,0.00013239411,0.00043731972,0.002445417,0.0077708773,0.008966893],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99787104,0.0006881691,0.0001254896,0.00045298453,0.00064148946,0.00022075062],"domain_scores_gemma":[0.9957911,0.0017138702,0.0001791622,0.0018734123,0.00035729608,0.00008520408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013845875,0.0011417358,0.0010061872,0.00092711905,0.0006295942,0.0013875734,0.0025463353,0.0015643679,0.0143075595],"category_scores_gemma":[0.0064490926,0.0007961433,0.0012443478,0.0007619802,0.0012423529,0.0017791417,0.0030459894,0.0016020774,0.005972639],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014761668,0.00022827575,0.0016362616,0.00084418035,0.00015785862,0.00056782056,0.0003417302,0.120533764,0.04313066,0.076253176,0.06288761,0.69194245],"study_design_scores_gemma":[0.00034198287,0.0003030014,0.00045429694,0.000064165564,0.00007922761,0.00040610522,0.00004455436,0.8160799,0.05654847,0.09891798,0.026670994,0.00008926765],"about_ca_topic_score_codex":0.0015554249,"about_ca_topic_score_gemma":0.0022930051,"teacher_disagreement_score":0.0143075595,"about_ca_system_score_codex":0.00067034987,"about_ca_system_score_gemma":0.0014107621,"threshold_uncertainty_score":0.047863543},"labels":[],"label_agreement":null},{"id":"W7117745108","doi":"10.1016/j.jss.2025.112748","title":"Exploring challenges in test mocking: Developer questions and insights from StackOverflow","year":2025,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Foundation for Innovation; University of Manitoba; University of Saskatchewan","keywords":"Popularity; Latent Dirichlet allocation; Topic model; Advice (programming); Key (lock); Selection (genetic algorithm); Test (biology)","score_opus":0.08815862776921252,"score_gpt":0.27019777538337264,"score_spread":0.18203914761416012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117745108","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9598111,0.0012497497,0.026868856,0.0066623227,0.000093845716,0.00012701516,0.0002314328,0.00030930998,0.004646314],"genre_scores_gemma":[0.9865639,0.0005012382,0.0098616965,0.00089366967,0.000077203185,0.000105763625,0.00027414848,0.00019282557,0.0015294879],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.97617984,0.015642554,0.0012149082,0.0018598718,0.0039526518,0.0011502855],"domain_scores_gemma":[0.7459812,0.21129747,0.016409885,0.0053578983,0.017117498,0.0038360762],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.033269923,0.0008072309,0.0005164175,0.004307127,0.0024698419,0.0043959497,0.0011643781,0.002085338,0.0009676348],"category_scores_gemma":[0.14246534,0.00070012576,0.00055380934,0.0019039996,0.0030168104,0.008576674,0.0038458966,0.002128894,0.00040700604],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024901197,0.0002180752,0.23627247,0.0007319797,0.00005660761,0.0021723362,0.59446377,0.0015428473,0.007740347,0.0052869897,0.007336909,0.14392874],"study_design_scores_gemma":[0.00005404975,0.00068795227,0.24193752,0.0014990249,0.00012506625,0.0035609428,0.60074276,0.024760481,0.008306187,0.016182851,0.101706125,0.00043705603],"about_ca_topic_score_codex":0.0029299413,"about_ca_topic_score_gemma":0.0055448823,"teacher_disagreement_score":0.96673006,"about_ca_system_score_codex":0.0028028535,"about_ca_system_score_gemma":0.0020212296,"threshold_uncertainty_score":0.17595029},"labels":[],"label_agreement":null},{"id":"W7120292292","doi":"","title":"Enhancing Knowledge Quality in Crowd-Sourced Developer Q&A Platforms through AI-driven Software Solutions","year":2025,"lang":"en","type":"article","venue":"University Library (University of Saskatchewan)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Code review; Leverage (statistics); Codebase; Code (set theory); Interpretability; Relevance (law); Quality (philosophy); Notice; Documentation","score_opus":0.01687023377960869,"score_gpt":0.2279958669204493,"score_spread":0.2111256331408406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7120292292","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4607558,0.0013740164,0.47927257,0.007059572,0.00031406904,0.002408052,0.0016828573,0.027595881,0.019537186],"genre_scores_gemma":[0.56337297,0.00052326306,0.42045555,0.0010533776,0.0001346819,0.0015275183,0.00327408,0.0026678808,0.0069906847],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9763613,0.012659886,0.0011948311,0.00373182,0.0052341097,0.00081805355],"domain_scores_gemma":[0.8506607,0.09993366,0.008271323,0.01984082,0.016477874,0.004815716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03389381,0.0014921643,0.0011058283,0.0058017513,0.0023199574,0.008101935,0.004122402,0.002685614,0.005657252],"category_scores_gemma":[0.13997844,0.0011371889,0.0011066645,0.0023258906,0.0024664244,0.008591161,0.011605826,0.0029392273,0.0032183272],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022185217,0.0032304814,0.03392262,0.0037471899,0.00038373694,0.002068417,0.074358486,0.02500168,0.035621822,0.021334033,0.032100674,0.76601225],"study_design_scores_gemma":[0.0018001368,0.0037834623,0.040314753,0.0028878432,0.000487668,0.0014072901,0.039107278,0.47723094,0.051025275,0.1288898,0.25207424,0.0009912619],"about_ca_topic_score_codex":0.0037039097,"about_ca_topic_score_gemma":0.004801584,"teacher_disagreement_score":0.03389381,"about_ca_system_score_codex":0.0027813984,"about_ca_system_score_gemma":0.0063607367,"threshold_uncertainty_score":0.17924976},"labels":[],"label_agreement":null},{"id":"W7123336555","doi":"10.1109/esem64174.2025.00052","title":"Exploring the Jupyter Ecosystem: An Empirical Study of Bugs and Vulnerabilities","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; Concordia University","funders":"","keywords":"Software deployment; Empirical research; Vulnerability (computing); Categorization; Software; Security bug; Grounded theory; Threat model","score_opus":0.11251177814678978,"score_gpt":0.3435553107639026,"score_spread":0.23104353261711283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123336555","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9976457,0.0001587854,0.00092439976,0.00019639736,0.000003873053,0.00007835614,0.0001538816,0.000038894785,0.00079961884],"genre_scores_gemma":[0.9941791,0.00037728067,0.004026788,0.00015222743,0.000012449234,0.0001701939,0.00042008216,0.000057142857,0.0006047992],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99533147,0.0017379947,0.00034047762,0.00072431506,0.0014240021,0.00044167496],"domain_scores_gemma":[0.9045075,0.0625586,0.018337157,0.0041423216,0.007329585,0.0031249013],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.008633617,0.0004932948,0.00038074944,0.0046086246,0.0019304766,0.002567917,0.0012968187,0.0010513529,0.0011628296],"category_scores_gemma":[0.053443145,0.0005666537,0.00025889697,0.003067749,0.0025583957,0.0060209157,0.003098807,0.0017369111,0.00026695136],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022396124,0.00083934906,0.6920155,0.0007710943,0.000053731528,0.0035200932,0.22697683,0.00058759574,0.0035827565,0.0018122913,0.0041151117,0.06550168],"study_design_scores_gemma":[0.000038490078,0.00070358854,0.77866983,0.0008402934,0.00006790032,0.0023879032,0.18648416,0.004404039,0.0024063701,0.0016913065,0.022186764,0.00011937739],"about_ca_topic_score_codex":0.0053541823,"about_ca_topic_score_gemma":0.009729684,"teacher_disagreement_score":0.9987032,"about_ca_system_score_codex":0.0019067643,"about_ca_system_score_gemma":0.002180517,"threshold_uncertainty_score":0.045659482},"labels":[],"label_agreement":null},{"id":"W7123340585","doi":"10.1109/esem64174.2025.00044","title":"Mapping Code Smells and Refactorings Accurately: Insights from an Empirical Study","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Code refactoring; Code smell; Code (set theory); Software; Software quality; Code review; Software maintenance; Source code","score_opus":0.08044891314345476,"score_gpt":0.3775742519305696,"score_spread":0.29712533878711483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123340585","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99252045,0.00061302213,0.0044530546,0.00042158965,0.000010606729,0.00009909332,0.00028069975,0.00005326651,0.0015481886],"genre_scores_gemma":[0.9944885,0.00031358525,0.004235826,0.00009549682,0.000009559172,0.00007629916,0.00038169886,0.000050167462,0.00034891334],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.96965516,0.013832291,0.0032695145,0.0025777395,0.009877778,0.0007875525],"domain_scores_gemma":[0.40212396,0.45198464,0.06703277,0.018160576,0.0584048,0.0022932566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034977674,0.00044544146,0.0004273413,0.0052180807,0.00090327294,0.0026599853,0.0014293758,0.0008755791,0.0010024518],"category_scores_gemma":[0.28251815,0.0005612869,0.0003958232,0.0055482937,0.001558962,0.004858169,0.0019004107,0.001669735,0.0003549681],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001781805,0.0005883695,0.8703915,0.00088888773,0.00015194798,0.0005757598,0.041581213,0.0012334924,0.0019750197,0.00075884955,0.0016580941,0.0800187],"study_design_scores_gemma":[0.00003735633,0.00057273585,0.944504,0.0007225588,0.00013176417,0.00065579324,0.030151831,0.010236134,0.003059444,0.0013094973,0.00852244,0.00009644837],"about_ca_topic_score_codex":0.008099049,"about_ca_topic_score_gemma":0.013722824,"teacher_disagreement_score":0.034977674,"about_ca_system_score_codex":0.0019506271,"about_ca_system_score_gemma":0.0024733841,"threshold_uncertainty_score":0.18498182},"labels":[],"label_agreement":null},{"id":"W7123342490","doi":"10.1109/esem64174.2025.00071","title":"What About Our Bug? A Study on the Responsiveness of NPM Package Maintainers","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Taxonomy (biology); Dependency (UML); Coding (social sciences); Software maintainer; Software; Best practice; Software bug","score_opus":0.02764741354663635,"score_gpt":0.32862160779007515,"score_spread":0.3009741942434388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123342490","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99679863,0.00020067157,0.0009962525,0.0006812453,0.000009594472,0.000042499654,0.000080971775,0.000025506079,0.0011646339],"genre_scores_gemma":[0.9984621,0.00015021485,0.0004975693,0.00027580958,0.000008915181,0.00006758096,0.00007440389,0.000024793408,0.0004387099],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9865715,0.006605529,0.0011586692,0.0014327653,0.0033823147,0.00084921415],"domain_scores_gemma":[0.8025376,0.12491502,0.046652805,0.004635012,0.017058346,0.004201286],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.020379394,0.00031541617,0.000414997,0.0037607355,0.0017730235,0.00280057,0.00112273,0.0010095994,0.0016950227],"category_scores_gemma":[0.13800807,0.0004367064,0.000377334,0.0025223382,0.0025026822,0.0046312166,0.003146193,0.0018178298,0.00035547587],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012378457,0.00014282126,0.53139365,0.000240832,0.000043459844,0.0006092216,0.42826936,0.000115959476,0.0014838453,0.00088877726,0.0014665623,0.035221774],"study_design_scores_gemma":[0.000013494696,0.0003632949,0.5863071,0.00031896267,0.00004086101,0.00067080214,0.4007933,0.0007001644,0.00069437124,0.00076102436,0.009259256,0.00007739134],"about_ca_topic_score_codex":0.005528318,"about_ca_topic_score_gemma":0.0056286957,"teacher_disagreement_score":0.9796206,"about_ca_system_score_codex":0.002773238,"about_ca_system_score_gemma":0.0020082395,"threshold_uncertainty_score":0.107777834},"labels":[],"label_agreement":null},{"id":"W7123349179","doi":"10.1109/esem64174.2025.00058","title":"A Fully Automated Agent for End-to-End Code Translation and Validation","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Correctness; Java; Code (set theory); Program comprehension; Source code; Static program analysis; Software; Code generation; Software development","score_opus":0.04136943109915722,"score_gpt":0.3246982982946076,"score_spread":0.28332886719545036,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123349179","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019880535,0.00006428839,0.9042389,0.00022905566,0.00008035367,0.00072655437,0.00023992559,0.071649715,0.0028906388],"genre_scores_gemma":[0.11866574,0.000060608745,0.8697848,0.00033388165,0.000035111854,0.0008274596,0.001028838,0.004571859,0.0046917247],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99603254,0.0015601837,0.00035732938,0.00069566764,0.0011634458,0.00019088207],"domain_scores_gemma":[0.98229086,0.007456358,0.0015913283,0.0047175703,0.0033792532,0.0005645417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046522324,0.0012110567,0.0006427795,0.0009641819,0.0006168678,0.0017517146,0.0020047524,0.0015553952,0.00795409],"category_scores_gemma":[0.0199112,0.0008026727,0.0008734347,0.00028374617,0.0010565324,0.0022220777,0.0032457798,0.0018534404,0.0060169357],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023917863,0.0017770858,0.011332276,0.0011913973,0.00028945322,0.001658212,0.0034363605,0.048832197,0.21357225,0.018645419,0.03991983,0.6569537],"study_design_scores_gemma":[0.00034157766,0.0007692673,0.0025421302,0.0002273771,0.0001086916,0.0008710138,0.00029938904,0.73505306,0.14194974,0.010519708,0.10714648,0.0001715236],"about_ca_topic_score_codex":0.00093700475,"about_ca_topic_score_gemma":0.0008521967,"teacher_disagreement_score":0.00795409,"about_ca_system_score_codex":0.00064360094,"about_ca_system_score_gemma":0.0018495128,"threshold_uncertainty_score":0.026609063},"labels":[],"label_agreement":null},{"id":"W7123350572","doi":"10.1109/esem64174.2025.00036","title":"Is LLM-Generated Code More Maintainable &amp; Reliable Than Human-Written Code?","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software quality; Maintainability; Python (programming language); Code review; Code (set theory); Software; Quality (philosophy); Key (lock)","score_opus":0.03301275349302725,"score_gpt":0.3218060359485615,"score_spread":0.28879328245553426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123350572","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89943117,0.001068091,0.083205044,0.0018620689,0.00012865092,0.00022076815,0.0013589496,0.0062866947,0.006438496],"genre_scores_gemma":[0.9379966,0.00022125678,0.057416458,0.00036679822,0.00003214581,0.00012257334,0.0018777667,0.0009203791,0.001046072],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9883117,0.005377899,0.00059936097,0.0019574952,0.0034258773,0.00032771652],"domain_scores_gemma":[0.8795105,0.075807795,0.01443164,0.017521206,0.0110566085,0.0016722508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012108054,0.00045933764,0.00039417803,0.0015668209,0.00039874035,0.00274059,0.0013793868,0.00094411825,0.001823263],"category_scores_gemma":[0.12867126,0.00036600168,0.0004997282,0.0009605999,0.0018029904,0.0023817625,0.0017143135,0.0009290808,0.00071012334],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001293112,0.00061677984,0.30771223,0.003653328,0.00045750593,0.0008332927,0.011059757,0.04270103,0.042911317,0.005225784,0.015863204,0.5676727],"study_design_scores_gemma":[0.00033152726,0.003047491,0.381717,0.0022773643,0.00046743875,0.0023038718,0.0077980096,0.4286526,0.06920277,0.025793679,0.07797721,0.0004311339],"about_ca_topic_score_codex":0.002369253,"about_ca_topic_score_gemma":0.0038242294,"teacher_disagreement_score":0.012108054,"about_ca_system_score_codex":0.0011221536,"about_ca_system_score_gemma":0.001662246,"threshold_uncertainty_score":0.06403428},"labels":[],"label_agreement":null},{"id":"W7123361151","doi":"10.1109/esem64174.2025.00064","title":"Understanding Everything as Code: A Taxonomy and Conceptual Model","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); Western University","funders":"","keywords":"Taxonomy (biology); Scope (computer science); Conceptual framework; Conceptual model; Scarcity; Work (physics)","score_opus":0.173478560142051,"score_gpt":0.3071178104514706,"score_spread":0.13363925030941956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123361151","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12506944,0.028896177,0.50393575,0.10345299,0.000969855,0.0025780045,0.0015991196,0.00068620755,0.23281245],"genre_scores_gemma":[0.67379403,0.02677745,0.27855387,0.0056059975,0.00027001667,0.0031387536,0.0023719163,0.00030328752,0.009184599],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9878449,0.0060172,0.0011524984,0.0014450104,0.0027611824,0.0007792144],"domain_scores_gemma":[0.9760272,0.014644239,0.0024249447,0.0014060037,0.0039209975,0.0015765051],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011665759,0.0014793533,0.0009942221,0.014976122,0.0069300616,0.0153873535,0.003978147,0.006051854,0.0035830508],"category_scores_gemma":[0.020255495,0.0008478775,0.0013528069,0.01855739,0.026864912,0.03909209,0.007281925,0.00562595,0.0008828797],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000010869823,0.000046830424,0.003735218,0.0006110621,0.000011406305,0.00025616004,0.09177918,0.00081015675,0.0003052071,0.8679485,0.003922309,0.030563068],"study_design_scores_gemma":[0.000012872226,0.000055643217,0.0031578175,0.0034695072,0.00002411635,0.0010094474,0.15879278,0.0058043585,0.00020041427,0.6227326,0.20466928,0.00007111929],"about_ca_topic_score_codex":0.020938233,"about_ca_topic_score_gemma":0.012233367,"teacher_disagreement_score":0.020938233,"about_ca_system_score_codex":0.013102921,"about_ca_system_score_gemma":0.020538233,"threshold_uncertainty_score":0.09506875},"labels":[],"label_agreement":null},{"id":"W7123361499","doi":"10.1109/esem64174.2025.00028","title":"Assessing Diversity in Creating Seed Set for Snowballing Search for Systematic Literature Review in Software Engineering","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Systematic review; Set (abstract data type); Diversity (politics); Software; Replication (statistics)","score_opus":0.043691024868300884,"score_gpt":0.3401598959469933,"score_spread":0.2964688710786924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123361499","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25193813,0.03257369,0.5347071,0.0065864935,0.0014704177,0.14705175,0.009241487,0.0034696797,0.012961296],"genre_scores_gemma":[0.24098651,0.002218728,0.6987794,0.0007714122,0.00017075114,0.054579016,0.001976329,0.00016781206,0.0003499784],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.546412,0.3007486,0.09502584,0.01884757,0.03674035,0.002225613],"domain_scores_gemma":[0.14938189,0.72139716,0.046379097,0.032516956,0.046965733,0.0033592323],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.44625896,0.0032567142,0.0061340965,0.056324452,0.0057561137,0.010371861,0.0044458043,0.0038397743,0.0037433268],"category_scores_gemma":[0.7710951,0.0026180495,0.009582103,0.029959485,0.004506833,0.010430793,0.013108132,0.0024720633,0.0010975406],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004799517,0.0007744953,0.10362995,0.15587223,0.020764621,0.001791206,0.04874599,0.015397576,0.0123872,0.020932931,0.014192265,0.60071206],"study_design_scores_gemma":[0.013036572,0.009950435,0.15669876,0.13653581,0.076360166,0.0044420087,0.026168453,0.150785,0.037724894,0.25017935,0.13441566,0.003702919],"about_ca_topic_score_codex":0.003352542,"about_ca_topic_score_gemma":0.010942411,"teacher_disagreement_score":0.55374104,"about_ca_system_score_codex":0.0076291375,"about_ca_system_score_gemma":0.03394559,"threshold_uncertainty_score":0.6828613},"labels":[],"label_agreement":null},{"id":"W7123362875","doi":"10.1109/esem64174.2025.00072","title":"How do Community Smells Influence Self-Admitted Technical Debt in Machine Learning Projects?","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Code smell; Technical debt; Quality (philosophy); Debt; Software quality","score_opus":0.016806396318620843,"score_gpt":0.2801697283172109,"score_spread":0.26336333199859,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123362875","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99734277,0.0001597215,0.00079652766,0.00040821082,0.000006635834,0.000014267571,0.00011258042,0.000021489073,0.0011378035],"genre_scores_gemma":[0.99928564,0.000055253506,0.00030727746,0.000024011959,0.0000097999255,0.000011901891,0.000084467625,0.000010186928,0.00021143901],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99366933,0.0023073456,0.0007833043,0.0009925577,0.0014805902,0.00076684216],"domain_scores_gemma":[0.83431137,0.055881593,0.08383746,0.005278672,0.010792549,0.009898317],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008045192,0.0002696128,0.00037437494,0.003160563,0.0011020472,0.003123463,0.0007704929,0.0009976204,0.002801216],"category_scores_gemma":[0.08334205,0.00035839135,0.0003346571,0.0033427447,0.0015838123,0.0046053557,0.0035034588,0.0011208353,0.0004968373],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007272346,0.000052228108,0.98165417,0.000043303742,0.000029738705,0.000116996845,0.004551361,0.00021028706,0.0002897067,0.00026025076,0.00039122757,0.012328032],"study_design_scores_gemma":[0.000004626325,0.000073050476,0.9893936,0.000045958004,0.000014024595,0.00013273871,0.007198996,0.0015148282,0.00015425048,0.0005722422,0.00087543996,0.000020369638],"about_ca_topic_score_codex":0.004757229,"about_ca_topic_score_gemma":0.008967269,"teacher_disagreement_score":0.9919548,"about_ca_system_score_codex":0.0013749041,"about_ca_system_score_gemma":0.0011080217,"threshold_uncertainty_score":0.042547584},"labels":[],"label_agreement":null},{"id":"W7124835252","doi":"10.1109/asew67777.2025.00075","title":"Beyond More Context: How Granularity and Order Drive Code Completion Quality","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Python (programming language); Context (archaeology); Code (set theory); Granularity; Source code; Quality (philosophy)","score_opus":0.029035892263101645,"score_gpt":0.323552250156846,"score_spread":0.29451635789374436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7124835252","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77379197,0.0040653716,0.19670911,0.0014214691,0.00019968908,0.00038007283,0.0013419802,0.016807735,0.00528262],"genre_scores_gemma":[0.92217976,0.00045436856,0.07275454,0.00020568103,0.0000576501,0.00011168338,0.0016359943,0.0011580321,0.0014422612],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9969494,0.000770957,0.00024875705,0.0008619196,0.00088255043,0.00028648585],"domain_scores_gemma":[0.9788831,0.012166239,0.0021176108,0.0030850705,0.0027518957,0.0009961075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039550615,0.0010189475,0.00089187303,0.0022059903,0.0006608337,0.0032398982,0.001415082,0.0012037503,0.0016250204],"category_scores_gemma":[0.052824903,0.0006570328,0.00072213373,0.0013786717,0.0010283451,0.0069419704,0.0022650948,0.0019958008,0.0011823722],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002858145,0.00090626976,0.13930653,0.001381327,0.00031100394,0.00056780997,0.002732864,0.18374906,0.0903784,0.005184195,0.015147204,0.5574772],"study_design_scores_gemma":[0.00019275781,0.0012420842,0.053567473,0.00017209364,0.00024212987,0.00046678275,0.001039145,0.8574134,0.06060426,0.014615814,0.010222901,0.00022119633],"about_ca_topic_score_codex":0.010077607,"about_ca_topic_score_gemma":0.0123061985,"teacher_disagreement_score":0.010077607,"about_ca_system_score_codex":0.0010678854,"about_ca_system_score_gemma":0.002188191,"threshold_uncertainty_score":0.02091664},"labels":[],"label_agreement":null},{"id":"W7124841212","doi":"10.1109/aiware69974.2025.00021","title":"CFCEval: Evaluating Security Aspects in Code Generated by Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Code (set theory); Metric (unit); Code review; Key (lock); Relevance (law)","score_opus":0.02988406004361717,"score_gpt":0.34952715752284097,"score_spread":0.3196430974792238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7124841212","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63573205,0.005049748,0.25726718,0.0012702339,0.0005541807,0.0011656597,0.019877542,0.06965945,0.009423875],"genre_scores_gemma":[0.6873966,0.000807642,0.25487053,0.0005682764,0.000076467135,0.0008098246,0.048812043,0.00428925,0.0023693533],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.990299,0.0039298194,0.0008914551,0.0015488248,0.0030583353,0.00027261683],"domain_scores_gemma":[0.9571205,0.029681504,0.0023857374,0.0050989445,0.0050575077,0.00065580284],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008362384,0.0017491715,0.0006293401,0.00369406,0.0005988063,0.0020325514,0.0021589973,0.0013880307,0.0013928342],"category_scores_gemma":[0.051455993,0.00042059325,0.001158228,0.002159027,0.0012564582,0.0030712336,0.0019486275,0.0017576002,0.00079877296],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022288996,0.0012660536,0.07100503,0.0043285904,0.0009976886,0.00068358,0.0017997572,0.27227405,0.03525683,0.01164979,0.07854353,0.5199662],"study_design_scores_gemma":[0.00028501812,0.001270846,0.014766843,0.00025572086,0.00015955699,0.00044008106,0.00039243774,0.91271454,0.039333105,0.010953388,0.019275416,0.00015296171],"about_ca_topic_score_codex":0.007102378,"about_ca_topic_score_gemma":0.012304413,"teacher_disagreement_score":0.008362384,"about_ca_system_score_codex":0.0017943708,"about_ca_system_score_gemma":0.0027190219,"threshold_uncertainty_score":0.044225037},"labels":[],"label_agreement":null},{"id":"W7124939946","doi":"10.1109/aiware69974.2025.00033","title":"Generative AI and Empirical Software Engineering: A Paradigm Shift","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Workflow; Paradigm shift; Generative grammar; Software; Software development; Social software engineering; Empirical research; Software evolution","score_opus":0.02193102504477759,"score_gpt":0.30693764596980755,"score_spread":0.28500662092502993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7124939946","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037481137,0.0799646,0.34595612,0.48006475,0.002859851,0.00021949771,0.00028821576,0.00039738172,0.052768484],"genre_scores_gemma":[0.6979624,0.04434932,0.20289148,0.04071609,0.0072707045,0.0012980623,0.00022105235,0.0006321625,0.004658774],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9284579,0.054342832,0.0022168695,0.004974285,0.00913145,0.0008765878],"domain_scores_gemma":[0.70696694,0.24526303,0.005262351,0.03041604,0.009258806,0.0028328837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08074897,0.0012264602,0.0024698302,0.011751886,0.0048343902,0.022474552,0.00501056,0.009948148,0.0035916565],"category_scores_gemma":[0.0872916,0.0015480916,0.0013429418,0.008709836,0.10152353,0.042025782,0.012597796,0.019805308,0.00080756546],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000014483623,0.0000437069,0.00050155417,0.00013979488,0.000013576112,0.00003003668,0.0028774103,0.0002565908,0.00009079933,0.98605853,0.0007376381,0.009235755],"study_design_scores_gemma":[0.000028865728,0.000028032002,0.00044300986,0.00034624198,0.0000069983025,0.00012357664,0.003009261,0.00169669,0.00013638282,0.9698238,0.024326593,0.000030709627],"about_ca_topic_score_codex":0.003399873,"about_ca_topic_score_gemma":0.0026374904,"teacher_disagreement_score":0.08074897,"about_ca_system_score_codex":0.0135841,"about_ca_system_score_gemma":0.011476713,"threshold_uncertainty_score":0.42704648},"labels":[],"label_agreement":null},{"id":"W7125579016","doi":"10.1109/cascon66301.2025.00073","title":"Workarounds in Software Forms: A User Study","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Hôtel-Dieu de Montréal","funders":"","keywords":"Workaround; Software; Software system; Software design; Software development","score_opus":0.018141091732122992,"score_gpt":0.30288414795126434,"score_spread":0.28474305621914137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125579016","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99558914,0.000054723227,0.0034066485,0.00010551528,0.00000515246,0.00013807377,0.00009767869,0.00014477792,0.0004583082],"genre_scores_gemma":[0.9911027,0.00011326293,0.007016521,0.00016699778,0.000008584704,0.00030458314,0.00018753624,0.000087446555,0.0010122529],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.98859847,0.00782378,0.00095602043,0.0008799029,0.001206066,0.00053582055],"domain_scores_gemma":[0.82154316,0.14453028,0.0061457655,0.011341663,0.014045444,0.0023937582],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01458286,0.0011882263,0.0008310818,0.0023245045,0.002068279,0.0021300754,0.0013520696,0.0017228067,0.0018289386],"category_scores_gemma":[0.06795381,0.001016534,0.00072906376,0.0011704403,0.0020765194,0.0034056294,0.0024004972,0.001489626,0.0006012414],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00274773,0.0077703884,0.28668553,0.0019239555,0.0002541304,0.0034424344,0.5329033,0.003716966,0.028859148,0.0025702708,0.0061095324,0.12301663],"study_design_scores_gemma":[0.0012418837,0.048215933,0.30715317,0.0017315382,0.0008605941,0.015615444,0.4266868,0.056078035,0.06946488,0.0048983432,0.06631481,0.001738588],"about_ca_topic_score_codex":0.002137068,"about_ca_topic_score_gemma":0.0030053358,"teacher_disagreement_score":0.01458286,"about_ca_system_score_codex":0.0011478667,"about_ca_system_score_gemma":0.0008884307,"threshold_uncertainty_score":0.07712245},"labels":[],"label_agreement":null},{"id":"W7125579974","doi":"10.1109/cascon66301.2025.00067","title":"Establishing Traceability Between Release Notes Heviand Software Artifacts: Practitioners' Perspectives","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University; University of Saskatchewan","funders":"","keywords":"Traceability; Requirements traceability; Usability; Software; Work (physics); Benchmark (surveying)","score_opus":0.02334393199472488,"score_gpt":0.30266354103632637,"score_spread":0.27931960904160147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125579974","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29109088,0.013026624,0.6496421,0.023384739,0.0003578487,0.0005032121,0.002750175,0.0075850487,0.011659326],"genre_scores_gemma":[0.5991931,0.0044292803,0.386128,0.0011508862,0.00018748858,0.00024231449,0.0056291176,0.0007685034,0.0022713323],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9518493,0.02180693,0.004491278,0.009504903,0.01129603,0.0010516619],"domain_scores_gemma":[0.71925145,0.15130626,0.02105832,0.047207188,0.057346344,0.003830361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05815439,0.0012386173,0.001128882,0.014417148,0.0015964813,0.009443732,0.0042401897,0.0025908821,0.001514733],"category_scores_gemma":[0.21911405,0.0010305184,0.00078176934,0.012179707,0.0027123918,0.014704373,0.005576251,0.0033157098,0.0011551089],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029973072,0.00047111465,0.17803241,0.0024082859,0.00022909355,0.000696536,0.01636914,0.013118215,0.012475461,0.018307667,0.013268779,0.74432355],"study_design_scores_gemma":[0.00016123925,0.0011257082,0.19275485,0.006027585,0.0004916033,0.0029369453,0.055915527,0.30068055,0.0626189,0.100003354,0.27672747,0.0005563112],"about_ca_topic_score_codex":0.012251011,"about_ca_topic_score_gemma":0.012025884,"teacher_disagreement_score":0.05815439,"about_ca_system_score_codex":0.002810884,"about_ca_system_score_gemma":0.005541683,"threshold_uncertainty_score":0.30755347},"labels":[],"label_agreement":null},{"id":"W7125580550","doi":"10.1109/cascon66301.2025.00069","title":"Can AI Build Systems? an Exploratory Study on Generating Software Architecture With LLMS","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Blueprint; Software architecture; Architecture; Software architecture description; Reference architecture; Software development; Resource-oriented architecture","score_opus":0.01849062765235339,"score_gpt":0.27648842184783623,"score_spread":0.25799779419548285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125580550","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9103428,0.00019063748,0.07447782,0.00108918,0.000019884497,0.00051782007,0.00019184191,0.00094548374,0.01222464],"genre_scores_gemma":[0.9145621,0.00014704315,0.08275456,0.00020555206,0.000007691266,0.0002236035,0.00027932294,0.00015149897,0.0016684622],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99642974,0.0027173434,0.000092500355,0.00026130877,0.00040459287,0.00009453414],"domain_scores_gemma":[0.9575946,0.036958836,0.0008925255,0.0030362003,0.0011826309,0.00033513305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006151986,0.00051963486,0.00025919112,0.00079722813,0.00072675204,0.0016593618,0.0011502879,0.0009571193,0.0026775198],"category_scores_gemma":[0.037524298,0.0004274525,0.00044594685,0.00064368866,0.0017589001,0.0046938187,0.0015287428,0.0015370285,0.00040003957],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017203507,0.0034376793,0.07584152,0.0021634556,0.00014291912,0.0023394786,0.10944078,0.20378342,0.039382838,0.13108487,0.0076412833,0.42302153],"study_design_scores_gemma":[0.00042436479,0.00401436,0.032681614,0.00044022335,0.00012794576,0.0010333452,0.04104831,0.7169549,0.034229036,0.09872664,0.07014732,0.00017201639],"about_ca_topic_score_codex":0.0030027411,"about_ca_topic_score_gemma":0.0053473986,"teacher_disagreement_score":0.006151986,"about_ca_system_score_codex":0.0013375074,"about_ca_system_score_gemma":0.000968243,"threshold_uncertainty_score":0.032535195},"labels":[],"label_agreement":null},{"id":"W7125581841","doi":"10.1109/cascon66301.2025.00075","title":"Understanding the Issue Types in Open Source Blockchain-Based Software Projects with the Transformer-Based BERTopic","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Software; Software development; Robustness (evolution); Open source; Open source software; Software evolution; Software peer review; Open-source software development","score_opus":0.05444498149886387,"score_gpt":0.288700487315476,"score_spread":0.23425550581661211,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125581841","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9743319,0.00068579084,0.02089404,0.0003693818,0.000010514516,0.00016070265,0.0010607094,0.00008531542,0.0024014844],"genre_scores_gemma":[0.98576295,0.00037789877,0.010745969,0.000044425407,0.000019038647,0.00015565686,0.0020515101,0.000047179474,0.0007954213],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9955923,0.0015793575,0.0004518988,0.00069188833,0.0012761165,0.00040838742],"domain_scores_gemma":[0.92498976,0.05372277,0.012732163,0.0023486225,0.005030957,0.0011758225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007997564,0.00031289758,0.00036000097,0.009464505,0.000863097,0.0025721365,0.00066888524,0.00073142396,0.0014102523],"category_scores_gemma":[0.051717285,0.0003607998,0.00047437247,0.008161118,0.0011451789,0.0067325933,0.0026954685,0.00093797565,0.00047198756],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036022792,0.00018986316,0.84508955,0.0009233046,0.00009640452,0.0006992477,0.028885929,0.005911119,0.0069397455,0.010709571,0.0022402294,0.09795481],"study_design_scores_gemma":[0.00004702516,0.00026032163,0.79607433,0.0004911623,0.000189362,0.0016207352,0.03747765,0.10883655,0.006114934,0.02507681,0.02371128,0.000099779536],"about_ca_topic_score_codex":0.004729772,"about_ca_topic_score_gemma":0.007543035,"teacher_disagreement_score":0.009464505,"about_ca_system_score_codex":0.0011973641,"about_ca_system_score_gemma":0.0014733698,"threshold_uncertainty_score":0.042295635},"labels":[],"label_agreement":null},{"id":"W7125585427","doi":"10.1109/cascon66301.2025.00071","title":"An Insight into the Technical Debt-Fix Trade Off in Software Backporting","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Technical debt; Eclipse; Python (programming language); Software; Debt; Software maintenance","score_opus":0.015418061659592096,"score_gpt":0.2987605363584914,"score_spread":0.28334247469889934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125585427","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95037633,0.0016806008,0.03029334,0.00322889,0.000072390634,0.0001154107,0.00043745662,0.00042326664,0.013372289],"genre_scores_gemma":[0.9893756,0.00032831906,0.0083031515,0.00018467878,0.000031015912,0.000031666063,0.00028879216,0.000100695026,0.0013560534],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.989484,0.0030375367,0.0010614883,0.0015949004,0.003721256,0.00110076],"domain_scores_gemma":[0.90525293,0.047935806,0.025410412,0.0086554475,0.00973828,0.0030070008],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01362875,0.00053200487,0.0004659282,0.0056316936,0.0017485264,0.004253696,0.001351523,0.0010595077,0.002706905],"category_scores_gemma":[0.08999286,0.0007881987,0.00053523487,0.005506534,0.002089468,0.0088193165,0.0035717902,0.0022441435,0.0005452104],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003434903,0.00028128485,0.7966969,0.00029779001,0.00014907836,0.002180862,0.011369337,0.008959322,0.006673348,0.020363985,0.0039650165,0.14871953],"study_design_scores_gemma":[0.000047368674,0.0005108387,0.8719209,0.000434193,0.00011251631,0.0040592207,0.01690077,0.041568637,0.004937681,0.038976733,0.020340813,0.0001902355],"about_ca_topic_score_codex":0.0059023723,"about_ca_topic_score_gemma":0.0073186303,"teacher_disagreement_score":0.01362875,"about_ca_system_score_codex":0.0026099265,"about_ca_system_score_gemma":0.0019922622,"threshold_uncertainty_score":0.07207656},"labels":[],"label_agreement":null},{"id":"W7125587998","doi":"10.1109/cascon66301.2025.00068","title":"ChatGPT for Code Refactoring: Analyzing Topics, Interaction, and Effective Prompts","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Code refactoring; Code (set theory); Software; Software maintenance; Source code; Software development","score_opus":0.020783934833576517,"score_gpt":0.3346795795495572,"score_spread":0.31389564471598064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125587998","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.65372175,0.0010687514,0.29918894,0.0018653546,0.00031760268,0.0022434853,0.010786634,0.02507171,0.0057357955],"genre_scores_gemma":[0.67623025,0.0004147831,0.30425337,0.0004698343,0.00012173942,0.0030975621,0.010181711,0.0014562494,0.0037745677],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9906144,0.0061227893,0.00052490545,0.0011472985,0.0013313738,0.0002592595],"domain_scores_gemma":[0.8497082,0.12895745,0.0063446397,0.006036033,0.0071418765,0.0018119015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009857169,0.0011636413,0.0008197512,0.0033365486,0.00078844355,0.0014118442,0.00099635,0.0009850613,0.0031407557],"category_scores_gemma":[0.090737,0.0004052053,0.00058805256,0.0022743272,0.00064611493,0.0035366893,0.002129604,0.0013383942,0.0016125399],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028603505,0.00071399636,0.10895651,0.0044773947,0.00018249206,0.0015161608,0.07782375,0.0049831546,0.056034703,0.0039312216,0.03613576,0.7023845],"study_design_scores_gemma":[0.0006833271,0.0037633672,0.3427213,0.0030425405,0.0005676576,0.0027567851,0.06256033,0.3455279,0.07294854,0.022270668,0.14232555,0.0008319949],"about_ca_topic_score_codex":0.0021755002,"about_ca_topic_score_gemma":0.004027125,"teacher_disagreement_score":0.009857169,"about_ca_system_score_codex":0.0010953378,"about_ca_system_score_gemma":0.0018970582,"threshold_uncertainty_score":0.052130282},"labels":[],"label_agreement":null},{"id":"W7125588785","doi":"10.1109/cascon66301.2025.00025","title":"Can We Trust the AI Pair Programmer? Copilot for API Misuse Detection and Correction","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Software; Precision and recall; Construct (python library); Coding (social sciences); Misuse detection; Secure coding; Code (set theory); Reliability (semiconductor)","score_opus":0.019992953017851853,"score_gpt":0.2875308711137979,"score_spread":0.267537918095946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125588785","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12121963,0.0008384601,0.6930993,0.007120342,0.00046483843,0.00043356433,0.000755014,0.15338877,0.022680128],"genre_scores_gemma":[0.45790935,0.00041053625,0.5062608,0.0020905295,0.000112788344,0.00026462882,0.0014132448,0.01960414,0.011933972],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98914796,0.003209416,0.0005309749,0.0016415433,0.0047534253,0.0007166299],"domain_scores_gemma":[0.95732576,0.014783099,0.0042124116,0.0143329315,0.007976017,0.0013698646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006276026,0.001166244,0.0007110206,0.0019734737,0.0010851548,0.003181431,0.002449529,0.0010442133,0.0057951314],"category_scores_gemma":[0.055064343,0.001315404,0.0007501781,0.0011560665,0.002178927,0.006602026,0.0037502327,0.003242007,0.0054979497],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011273331,0.00041675646,0.040761285,0.0007996044,0.00013465885,0.0017591699,0.00892724,0.007325319,0.04154293,0.021060737,0.094380595,0.7817643],"study_design_scores_gemma":[0.00022507731,0.0008146154,0.024933897,0.0010621842,0.00026092923,0.0069603734,0.0033413502,0.39171952,0.13768218,0.0356745,0.3969216,0.00040379335],"about_ca_topic_score_codex":0.0032129528,"about_ca_topic_score_gemma":0.0040113237,"teacher_disagreement_score":0.006276026,"about_ca_system_score_codex":0.0014432624,"about_ca_system_score_gemma":0.0030184013,"threshold_uncertainty_score":0.033191204},"labels":[],"label_agreement":null},{"id":"W7125591383","doi":"10.1109/cascon66301.2025.00110","title":"Different by Design: Understanding Human-AI Collaboration Through GitHub Pull Requests","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Maintainability; Coding (social sciences); Software; Bridge (graph theory); Cursor (databases); Descriptive statistics; Statistical analysis; Presentation (obstetrics)","score_opus":0.07299434836039533,"score_gpt":0.3513910749916013,"score_spread":0.278396726631206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125591383","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84814316,0.0015000597,0.10779927,0.0030173801,0.00012055336,0.00033790377,0.010509517,0.0064932923,0.02207899],"genre_scores_gemma":[0.9206562,0.00036034812,0.056827858,0.00043652407,0.000046489484,0.00037942655,0.016445348,0.0010277329,0.0038199327],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99551165,0.0025853538,0.00022113447,0.00066910096,0.0007571968,0.0002557337],"domain_scores_gemma":[0.96648586,0.02118169,0.0034742355,0.0056049908,0.0024602304,0.00079293014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053438516,0.0005337231,0.00030347178,0.002894834,0.000921679,0.0035164994,0.001496669,0.001163709,0.0016776546],"category_scores_gemma":[0.038637992,0.0003001882,0.00041023665,0.0022443656,0.0011516112,0.0049404483,0.0026817394,0.0010044369,0.0012500766],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013683491,0.000874168,0.48333022,0.0017581449,0.00023842399,0.0012688927,0.058326412,0.017095653,0.011839918,0.03858154,0.057690762,0.32762745],"study_design_scores_gemma":[0.00014093776,0.000625574,0.27475035,0.0007440293,0.0001500143,0.0013947703,0.042725973,0.30381712,0.011370856,0.08034876,0.28368896,0.00024255637],"about_ca_topic_score_codex":0.007179373,"about_ca_topic_score_gemma":0.014227849,"teacher_disagreement_score":0.007179373,"about_ca_system_score_codex":0.0012351867,"about_ca_system_score_gemma":0.0010153987,"threshold_uncertainty_score":0.028261364},"labels":[],"label_agreement":null},{"id":"W7125592731","doi":"10.1109/cascon66301.2025.00101","title":"LLM-Powered Code Quality Bot: Automating Refactoring in CI/CD","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Code refactoring; Code (set theory); Code smell; Pipeline (software); Focus (optics); Software; Static program analysis; Software quality; Quality (philosophy)","score_opus":0.04801761847803115,"score_gpt":0.37377263786362347,"score_spread":0.3257550193855923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125592731","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05449148,0.00073682086,0.5407014,0.0011044911,0.00037535295,0.0006363498,0.0011606968,0.38711795,0.013675464],"genre_scores_gemma":[0.22508898,0.00029545432,0.73711103,0.00077464565,0.00006916563,0.00032898036,0.00198767,0.020788204,0.013555863],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.996729,0.0006465733,0.00017251946,0.000766657,0.0014490491,0.00023617006],"domain_scores_gemma":[0.9863391,0.0056062113,0.0013612021,0.0035645645,0.002407208,0.000721623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004462159,0.0011838045,0.00056376075,0.0027636925,0.0009110729,0.0019159281,0.0019006894,0.0012709863,0.009056145],"category_scores_gemma":[0.02127101,0.0009847995,0.00065163017,0.00086850347,0.00094614306,0.0028103723,0.0030481,0.0020458126,0.0055731214],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087997573,0.00033402737,0.011296353,0.0008507031,0.00009066269,0.00053189,0.0028021592,0.005840424,0.088311955,0.0077361874,0.09333099,0.78799474],"study_design_scores_gemma":[0.00042697752,0.0007823459,0.013989577,0.00046024492,0.00016314867,0.0015961011,0.001218546,0.45069852,0.20544156,0.01581657,0.30900586,0.00040064644],"about_ca_topic_score_codex":0.0047656572,"about_ca_topic_score_gemma":0.007902344,"teacher_disagreement_score":0.009056145,"about_ca_system_score_codex":0.0016199595,"about_ca_system_score_gemma":0.0031945338,"threshold_uncertainty_score":0.030295849},"labels":[],"label_agreement":null},{"id":"W7125598063","doi":"10.1109/cascon66301.2025.00045","title":"An Empirical Evaluation of LLM-Based Approaches for Code Vulnerability Detection: RAG, SFT, and Dual-Agent Systems","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Transparency (behavior); Vulnerability (computing); Audit; Baseline (sea); Domain (mathematical analysis); Software; Task (project management); The Internet","score_opus":0.1558992302502042,"score_gpt":0.3882194443479205,"score_spread":0.2323202140977163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125598063","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88434553,0.0054506743,0.084004276,0.0010597737,0.00024681597,0.00086440716,0.003675642,0.0147621175,0.005590802],"genre_scores_gemma":[0.8687012,0.0005937717,0.12058058,0.00023115684,0.000077393444,0.00033828834,0.0075522196,0.00043426643,0.0014911284],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9864417,0.0066787456,0.0011597394,0.0026162514,0.0026252489,0.0004783937],"domain_scores_gemma":[0.9553109,0.030954547,0.0025471915,0.0068026786,0.0033209713,0.0010636816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014832025,0.0017865747,0.0009904442,0.0062366207,0.0008842111,0.0016477274,0.0025794832,0.0021751297,0.0012144019],"category_scores_gemma":[0.04701111,0.0004710242,0.0011794543,0.0026216889,0.0014165761,0.004048508,0.003848963,0.0020064665,0.0010003805],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036312318,0.003287518,0.12769435,0.0031204715,0.0013509765,0.0004717375,0.0022927658,0.13850126,0.010082789,0.002999798,0.023089508,0.6834776],"study_design_scores_gemma":[0.00028325507,0.0019102378,0.029708354,0.00024152976,0.00036172534,0.0005938048,0.0009280176,0.9386735,0.014757458,0.0040793456,0.008329694,0.00013299342],"about_ca_topic_score_codex":0.00665756,"about_ca_topic_score_gemma":0.007853733,"teacher_disagreement_score":0.014832025,"about_ca_system_score_codex":0.0019228214,"about_ca_system_score_gemma":0.0019944725,"threshold_uncertainty_score":0.07844013},"labels":[],"label_agreement":null},{"id":"W7125603952","doi":"10.1109/cascon66301.2025.00078","title":"Adaptive Fine-Tuning for Multi-Class Classification Over Software Requirement Data","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Software requirements; Software; Robustness (evolution); Transfer of learning; Pooling; Task (project management); Field (mathematics); Deep learning; Software development","score_opus":0.20733667833288394,"score_gpt":0.3838192928120084,"score_spread":0.17648261447912444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125603952","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24671611,0.0016097205,0.7333574,0.0009302347,0.00030406934,0.00034159602,0.0011507358,0.013755835,0.0018342327],"genre_scores_gemma":[0.82068586,0.00023505361,0.17054343,0.0008505408,0.00016101962,0.00041234744,0.003874076,0.00039631833,0.0028412968],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982849,0.0004701443,0.00012067708,0.00074407144,0.00018265992,0.00019759552],"domain_scores_gemma":[0.9957858,0.0024564588,0.00029840125,0.00070350536,0.0005678945,0.00018790699],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030649053,0.0014303277,0.0013776185,0.0013647577,0.00052639656,0.0008381496,0.0023312643,0.0018349924,0.0017833242],"category_scores_gemma":[0.008580158,0.00046644305,0.001185312,0.0010633786,0.0007257115,0.0022612142,0.0013976356,0.0031240745,0.0013318298],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084132113,0.001417433,0.010263245,0.0003036675,0.00021148601,0.00026665564,0.00037175085,0.2502007,0.030841194,0.0014400234,0.013286082,0.69055647],"study_design_scores_gemma":[0.000022577458,0.00006185448,0.0009314935,0.000007469087,0.000011058447,0.000024523473,0.000040494153,0.9937571,0.0028824403,0.0018200687,0.00042935414,0.000011537212],"about_ca_topic_score_codex":0.007086873,"about_ca_topic_score_gemma":0.009429834,"teacher_disagreement_score":0.007086873,"about_ca_system_score_codex":0.00095971674,"about_ca_system_score_gemma":0.0012994114,"threshold_uncertainty_score":0.016208947},"labels":[],"label_agreement":null},{"id":"W7125610271","doi":"10.1109/cascon66301.2025.00076","title":"Multi-Label Ambiguity Detection in Software Requirements Using Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Vagueness; Ambiguity; Software requirements specification; Scope (computer science); Software requirements; Natural language; Set (abstract data type); Software; Focus (optics)","score_opus":0.08411771468295513,"score_gpt":0.3581302819086264,"score_spread":0.27401256722567124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125610271","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5098052,0.0027878673,0.45800373,0.0014424723,0.00031399002,0.0005651761,0.0054926137,0.017172106,0.004416791],"genre_scores_gemma":[0.72302884,0.00040061658,0.25440347,0.00051551766,0.00009989107,0.0005103248,0.018193185,0.00051920576,0.002328917],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99219525,0.003607072,0.0006598117,0.001832419,0.0013779601,0.00032760695],"domain_scores_gemma":[0.9815726,0.013062956,0.001136032,0.0017050665,0.002234283,0.00028898538],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006496515,0.0017922386,0.000951062,0.003916373,0.00093786354,0.0016709825,0.0016328807,0.0017717447,0.0007569188],"category_scores_gemma":[0.018564709,0.00049104483,0.002020472,0.0016996824,0.00063326664,0.0024093173,0.00208771,0.0021167335,0.00090279634],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016651521,0.0014386683,0.041823708,0.0017129736,0.00042867038,0.0018665065,0.0022547576,0.13155048,0.048734326,0.0050346083,0.030763965,0.7327262],"study_design_scores_gemma":[0.00011053822,0.00029852605,0.010497536,0.000117920754,0.00013739502,0.0005345848,0.0006496026,0.9457651,0.02227615,0.007168153,0.012348323,0.00009630655],"about_ca_topic_score_codex":0.0074026915,"about_ca_topic_score_gemma":0.0126862535,"teacher_disagreement_score":0.0074026915,"about_ca_system_score_codex":0.0014429297,"about_ca_system_score_gemma":0.0015293871,"threshold_uncertainty_score":0.03435731},"labels":[],"label_agreement":null},{"id":"W7125648772","doi":"10.1109/cascon66301.2025.00080","title":"Reverse Engineering User Stories from Code using Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Reverse engineering; Code (set theory); Source code; Modeling language; Natural language; User interface","score_opus":0.024737771478988667,"score_gpt":0.29032718384310374,"score_spread":0.2655894123641151,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125648772","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39263025,0.0006044714,0.5580548,0.0048445393,0.0002859801,0.0006422255,0.005486568,0.01003076,0.027420498],"genre_scores_gemma":[0.77877736,0.0004696585,0.19768262,0.00067832775,0.00007437766,0.00030188222,0.0058424952,0.0035126521,0.012660577],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9961234,0.0019614466,0.0001368617,0.00032205333,0.0013136624,0.00014252268],"domain_scores_gemma":[0.9412203,0.04689437,0.0017804637,0.004425503,0.005340095,0.00033940113],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024390172,0.00090903736,0.00034709607,0.0016708593,0.0008615958,0.0022103724,0.0011904383,0.0012442239,0.0051313667],"category_scores_gemma":[0.043040432,0.0006103762,0.00067403953,0.00095119927,0.0010981419,0.0044656442,0.0020125199,0.0021627285,0.0017295633],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017773996,0.0013166431,0.0333884,0.0029018293,0.00028593728,0.009648502,0.074635215,0.03696661,0.06794515,0.08288073,0.06874869,0.619505],"study_design_scores_gemma":[0.00018596351,0.000723796,0.013028429,0.0012356645,0.00041985733,0.0067614294,0.020878643,0.4704447,0.13297032,0.08100915,0.27201173,0.00033025516],"about_ca_topic_score_codex":0.0042203707,"about_ca_topic_score_gemma":0.008705552,"teacher_disagreement_score":0.0051313667,"about_ca_system_score_codex":0.0007563218,"about_ca_system_score_gemma":0.0012171263,"threshold_uncertainty_score":0.017166138},"labels":[],"label_agreement":null},{"id":"W7125651587","doi":"10.1109/cascon66301.2025.00065","title":"Design, Implementation and Evaluation of a Novel Programming Language Topic Classification Workflow","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Workflow; Natural language; Feature (linguistics); Semantics (computer science)","score_opus":0.10155059606357696,"score_gpt":0.40851198459399946,"score_spread":0.3069613885304225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125651587","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025948307,0.00022593637,0.8631652,0.0006138855,0.00021754563,0.0021497027,0.0010403733,0.10471034,0.0019286844],"genre_scores_gemma":[0.09780779,0.00017251294,0.8879589,0.00036789002,0.00005468588,0.0010161054,0.0028458517,0.0058380584,0.003938199],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946373,0.001261196,0.00076251186,0.001252764,0.0017341818,0.00035199212],"domain_scores_gemma":[0.9809585,0.00868226,0.00078737165,0.003674383,0.004630998,0.0012664876],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009836649,0.00108847,0.0010016031,0.0020088244,0.0011610166,0.0058063236,0.004155999,0.0013083545,0.007859108],"category_scores_gemma":[0.02590121,0.001285224,0.0012835665,0.0014684526,0.0010630434,0.005351818,0.0030098166,0.0026552482,0.004644043],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003948678,0.0024543165,0.011170079,0.0016280874,0.00051426433,0.00043652122,0.0028267875,0.011384579,0.11646652,0.0212242,0.038811352,0.78913456],"study_design_scores_gemma":[0.0018653788,0.0018122907,0.0059228414,0.00029428533,0.00040504037,0.0004510419,0.0011374837,0.5888088,0.27037084,0.01921896,0.10938329,0.00032973516],"about_ca_topic_score_codex":0.008984847,"about_ca_topic_score_gemma":0.00793487,"teacher_disagreement_score":0.009836649,"about_ca_system_score_codex":0.0017742383,"about_ca_system_score_gemma":0.0063161566,"threshold_uncertainty_score":0.052021742},"labels":[],"label_agreement":null},{"id":"W7125655103","doi":"10.1109/cascon66301.2025.00092","title":"RevMine: An LLM-Assisted Tool for Code Review Mining and Analysis Across Git Platforms","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Stroke Network; École de Technologie Supérieure","funders":"","keywords":"Code (set theory); Source code; Key (lock); Identification (biology)","score_opus":0.044368592443285365,"score_gpt":0.3736740081863399,"score_spread":0.3293054157430545,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125655103","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017469611,0.0010861348,0.32150856,0.0010782551,0.00025624706,0.0008237905,0.045617383,0.60513633,0.0070237448],"genre_scores_gemma":[0.09120782,0.0007632235,0.7538903,0.0008982634,0.00018051869,0.0015825176,0.09458575,0.04410478,0.012786879],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99621576,0.00085887214,0.00059943047,0.00082753913,0.0012928648,0.00020551906],"domain_scores_gemma":[0.97755593,0.011951238,0.0035336907,0.0032461667,0.0029502516,0.00076265755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005319552,0.0019387401,0.00093980425,0.010760725,0.001287728,0.003618966,0.0028019075,0.0013417662,0.019448856],"category_scores_gemma":[0.03284646,0.0011421922,0.0021775495,0.0038342753,0.00077411556,0.00380534,0.0037991144,0.0021135863,0.011142185],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075823104,0.0003475012,0.026144182,0.0066392897,0.0008805739,0.0009938043,0.0023391338,0.012102781,0.021895276,0.016470851,0.36842784,0.5430005],"study_design_scores_gemma":[0.00053237757,0.0004725392,0.02152684,0.0014778889,0.00059336715,0.0015399932,0.0014251864,0.375301,0.06372448,0.04058188,0.49242446,0.0003999228],"about_ca_topic_score_codex":0.0054465723,"about_ca_topic_score_gemma":0.015393335,"teacher_disagreement_score":0.019448856,"about_ca_system_score_codex":0.0011896867,"about_ca_system_score_gemma":0.0051984005,"threshold_uncertainty_score":0.06506288},"labels":[],"label_agreement":null},{"id":"W7125948450","doi":"10.1109/ase63991.2025.00077","title":"An Empirical Study of Python Library Migration Using Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Python (programming language); Code (set theory); Benchmark (surveying); Empirical research; Unit testing","score_opus":0.03883258264504589,"score_gpt":0.36056596619006676,"score_spread":0.32173338354502085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125948450","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99448806,0.00018355109,0.0026443824,0.00024479325,0.000019573441,0.000089207395,0.00066372764,0.00067704805,0.0009896363],"genre_scores_gemma":[0.98660815,0.00013241047,0.009088534,0.00011510034,0.0000106308025,0.00014692819,0.0031473527,0.00024782232,0.00050311873],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98847896,0.006101674,0.00095237524,0.0013483479,0.0026544356,0.000464188],"domain_scores_gemma":[0.8809332,0.085512295,0.01063594,0.012581767,0.008475722,0.0018610755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014504657,0.0008777505,0.0004760058,0.0016070551,0.00084999047,0.0015798693,0.0018636821,0.000965164,0.0008251935],"category_scores_gemma":[0.11290673,0.0006112765,0.0007183683,0.0026096832,0.0015850503,0.004491808,0.0016456657,0.0020780268,0.00045831877],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017152131,0.003508136,0.78829116,0.0012409444,0.0004950311,0.0011364995,0.007974189,0.0699876,0.006160842,0.0035611968,0.0135612395,0.102368064],"study_design_scores_gemma":[0.00037969777,0.00296838,0.3632089,0.00033683723,0.00034432282,0.0014611833,0.0077737705,0.5864141,0.010965311,0.004490262,0.021445833,0.00021136175],"about_ca_topic_score_codex":0.011170519,"about_ca_topic_score_gemma":0.013428964,"teacher_disagreement_score":0.014504657,"about_ca_system_score_codex":0.0018468825,"about_ca_system_score_gemma":0.0014960584,"threshold_uncertainty_score":0.07670885},"labels":[],"label_agreement":null},{"id":"W7125948853","doi":"10.1109/ase63991.2025.00086","title":"Coverage-Based Harmfulness Testing for LLM Code Transformation","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Moderation; Code (set theory); Set (abstract data type); Source code; Harm; Key (lock); Software; Code review","score_opus":0.04030988784248639,"score_gpt":0.3051761466031735,"score_spread":0.2648662587606871,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125948853","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36173233,0.00084344944,0.59174126,0.0006462155,0.000086694025,0.00096941384,0.0009922122,0.0389977,0.0039907694],"genre_scores_gemma":[0.7383638,0.00013192891,0.25749415,0.00027611863,0.00001986373,0.00040111705,0.0010267796,0.0013126021,0.0009735443],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9881187,0.0034964005,0.0006724148,0.002101157,0.0050254865,0.0005859187],"domain_scores_gemma":[0.9472181,0.035296608,0.0045503876,0.006975554,0.005186506,0.0007729405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004898579,0.0016798243,0.0008261545,0.0020739578,0.00061810366,0.0011082861,0.0022836525,0.0013836761,0.0019150722],"category_scores_gemma":[0.051947247,0.0007079,0.0012090931,0.00066949654,0.0022382203,0.0028038465,0.0022961863,0.0017379748,0.0006529441],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021731309,0.0017181424,0.08573601,0.0018426912,0.00033554077,0.0013205012,0.0033179407,0.14670937,0.19272307,0.012071367,0.007255949,0.5447963],"study_design_scores_gemma":[0.00015417735,0.0013119568,0.013209894,0.0001779268,0.00015747949,0.0007755693,0.00041539926,0.8695294,0.0974845,0.011364906,0.0053024814,0.000116283096],"about_ca_topic_score_codex":0.0047554118,"about_ca_topic_score_gemma":0.0062586227,"teacher_disagreement_score":0.004898579,"about_ca_system_score_codex":0.0012956804,"about_ca_system_score_gemma":0.0029153216,"threshold_uncertainty_score":0.025906444},"labels":[],"label_agreement":null},{"id":"W7125949815","doi":"10.1109/ase63991.2025.00192","title":"SPICE: An Automated SWE-Bench Labeling Pipeline for Issue Clarity, Test Coverage, and Effort Estimation","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Queen's University","funders":"","keywords":"Spice; Pipeline (software); Software; Test (biology); Code (set theory); Test data","score_opus":0.015864431493693845,"score_gpt":0.3423301053617524,"score_spread":0.3264656738680586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125949815","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09514207,0.0025036498,0.35676518,0.0029750564,0.00078341685,0.0012965424,0.318856,0.19410759,0.027570566],"genre_scores_gemma":[0.11544639,0.00047979705,0.3508419,0.0010927813,0.00016566768,0.0018528748,0.5142293,0.008032863,0.007858393],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995747,0.0008535919,0.00044366234,0.0014912427,0.0011851502,0.00027933763],"domain_scores_gemma":[0.98002213,0.008152458,0.0016961019,0.0048793447,0.0046538366,0.00059619005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004157213,0.0020205826,0.0008738875,0.006329111,0.001389264,0.0022744245,0.0023282985,0.0021008132,0.007965272],"category_scores_gemma":[0.027535671,0.000720125,0.0015476752,0.0035325529,0.00077451,0.0033255243,0.003642183,0.0024938982,0.007392619],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006347236,0.0004701443,0.038143072,0.0030890238,0.00023081154,0.00064873544,0.0019298667,0.013346094,0.027841704,0.010397584,0.6043104,0.29895774],"study_design_scores_gemma":[0.00040079135,0.00040043684,0.05679619,0.0008627982,0.00020019693,0.00074876123,0.0018209227,0.18635954,0.05074955,0.028998712,0.6724016,0.00026057567],"about_ca_topic_score_codex":0.007033782,"about_ca_topic_score_gemma":0.026008578,"teacher_disagreement_score":0.007965272,"about_ca_system_score_codex":0.0014019294,"about_ca_system_score_gemma":0.00287264,"threshold_uncertainty_score":0.026646495},"labels":[],"label_agreement":null},{"id":"W7125963704","doi":"10.1109/ase63991.2025.00271","title":"Context-Aware CodeLLM Eviction for AI-assisted Coding","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Western University; Huawei Technologies (Canada)","funders":"","keywords":"Eviction; Coding (social sciences); Latency (audio); Software; Code (set theory); Resource (disambiguation)","score_opus":0.037967209368414113,"score_gpt":0.32664985764502624,"score_spread":0.2886826482766121,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125963704","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20889243,0.0012309087,0.75763613,0.0006513649,0.00019884726,0.00017313672,0.00017886775,0.026836857,0.004201481],"genre_scores_gemma":[0.7990087,0.00022814119,0.19679919,0.0002989023,0.00003593663,0.00009898605,0.00028216673,0.0011216793,0.0021263731],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986093,0.00027894942,0.00010346654,0.00023207643,0.00054797047,0.00022819574],"domain_scores_gemma":[0.99529123,0.001563763,0.00042511473,0.001769856,0.00067896675,0.00027099618],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011505027,0.00074517285,0.00039918476,0.000696227,0.00060778565,0.00093882467,0.0019401421,0.0005339549,0.0013469864],"category_scores_gemma":[0.008479268,0.0003885651,0.00035876877,0.0005773547,0.00066016486,0.0023285754,0.0026446423,0.0016014994,0.0005605347],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001103069,0.0005616546,0.014507495,0.0003143059,0.00009609401,0.000829203,0.0019827627,0.15535665,0.16136889,0.02053709,0.015405682,0.62793714],"study_design_scores_gemma":[0.000044954013,0.00014271811,0.0012212363,0.000031224226,0.000028431554,0.00017983337,0.00021681972,0.91855085,0.061373938,0.008156634,0.010009908,0.000043392945],"about_ca_topic_score_codex":0.0039041482,"about_ca_topic_score_gemma":0.00862161,"teacher_disagreement_score":0.0039041482,"about_ca_system_score_codex":0.00096613367,"about_ca_system_score_gemma":0.0022964769,"threshold_uncertainty_score":0.0077627897},"labels":[],"label_agreement":null},{"id":"W7125971062","doi":"10.1109/ase63991.2025.00201","title":"Token Sugar: Making Source Code Sweeter for LLMs through Token-Efficient Shorthand","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Google","keywords":"Security token; Source code; Boilerplate text; Code (set theory); Code generation; Inference; Language model","score_opus":0.03464808555862875,"score_gpt":0.3268847249025322,"score_spread":0.2922366393439034,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125971062","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09964423,0.0015652941,0.7026022,0.0010121547,0.0008304719,0.00043226013,0.0046258555,0.18498589,0.004301629],"genre_scores_gemma":[0.311627,0.0005078228,0.64173985,0.0014113602,0.00014153303,0.0007577651,0.022691535,0.0086626895,0.012460466],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99874496,0.00025727684,0.00008705145,0.00056714506,0.00022568059,0.0001178543],"domain_scores_gemma":[0.9968363,0.001443536,0.00023752482,0.000889542,0.00043500043,0.00015802839],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017520915,0.0027656567,0.00093437434,0.0019791874,0.0006424604,0.0015739526,0.0035908583,0.0017056429,0.0064360104],"category_scores_gemma":[0.0077260095,0.0010982335,0.0018495752,0.00113603,0.00115688,0.0040373835,0.0030970338,0.0040801195,0.0062986854],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090129924,0.00039524015,0.008011831,0.00073568174,0.00027476408,0.0006171765,0.0006532767,0.07437397,0.035678238,0.0060710306,0.053541947,0.81874555],"study_design_scores_gemma":[0.0001265127,0.00029241786,0.0012682733,0.00008405521,0.000118226606,0.00023759075,0.0002114561,0.9232066,0.039447077,0.014440571,0.020487398,0.00007981187],"about_ca_topic_score_codex":0.0060300715,"about_ca_topic_score_gemma":0.015488379,"teacher_disagreement_score":0.0064360104,"about_ca_system_score_codex":0.0009567304,"about_ca_system_score_gemma":0.0023724742,"threshold_uncertainty_score":0.021530628},"labels":[],"label_agreement":null},{"id":"W7125980810","doi":"10.1109/ase63991.2025.00241","title":"The Cost of Downgrading Build Systems : A Case Study of Kubernetes","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Waterloo","funders":"","keywords":"Downgrade; Maintainability; Work (physics); Resource (disambiguation); Technical debt; Undo; Container (type theory)","score_opus":0.023816464824489163,"score_gpt":0.3088211013379941,"score_spread":0.28500463651350494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125980810","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99306196,0.0002612298,0.0033861322,0.00024431985,0.000024519644,0.00007431811,0.00023422539,0.0006865457,0.0020266671],"genre_scores_gemma":[0.9869082,0.00019372188,0.010619169,0.00008149639,0.000013785515,0.000067172645,0.00064830127,0.00038274608,0.0010854073],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98908037,0.0037452208,0.0008984861,0.0010622131,0.004064329,0.0011494345],"domain_scores_gemma":[0.94077367,0.032606285,0.0063419966,0.009708396,0.0077626444,0.002806988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077966726,0.0010017877,0.00053205807,0.0033167452,0.0015603587,0.002141231,0.0023884282,0.0008909866,0.0016528497],"category_scores_gemma":[0.049136918,0.0008677685,0.0007122508,0.0033517964,0.0020461807,0.0034277483,0.002373139,0.0026564344,0.00052403036],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036211424,0.002717432,0.4680493,0.0018360614,0.00040488673,0.011249881,0.020835951,0.07217987,0.034756493,0.011889492,0.011204721,0.36125487],"study_design_scores_gemma":[0.00057876814,0.0059659947,0.6560659,0.0008431819,0.0007823388,0.009512295,0.02592065,0.16447173,0.059729036,0.011713904,0.06371721,0.00069907703],"about_ca_topic_score_codex":0.011126074,"about_ca_topic_score_gemma":0.014219285,"teacher_disagreement_score":0.011126074,"about_ca_system_score_codex":0.0025017764,"about_ca_system_score_gemma":0.0018068313,"threshold_uncertainty_score":0.04123324},"labels":[],"label_agreement":null},{"id":"W7125984020","doi":"10.1109/ase63991.2025.00137","title":"Characterizing Multi-Hunk Patches: Divergence, Proximity, and LLM Repair Challenges","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of British Columbia","funders":"","keywords":"Metric (unit); Disjoint sets; Divergence (linguistics); Empirical research; Variation (astronomy); Event (particle physics)","score_opus":0.04611034680193843,"score_gpt":0.2915173091313208,"score_spread":0.24540696232938236,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125984020","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.87480515,0.0027475082,0.110931695,0.0009782866,0.00008398095,0.00015450601,0.0036368503,0.004290432,0.0023716309],"genre_scores_gemma":[0.946671,0.00035325132,0.044641014,0.00017488343,0.000035561217,0.00011283217,0.006751755,0.00054056133,0.00071913085],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9942966,0.0012974769,0.00043713427,0.001998556,0.0016156109,0.00035465328],"domain_scores_gemma":[0.95096254,0.027402882,0.0065438524,0.010720915,0.003162035,0.0012077892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064547565,0.0009753499,0.0010289826,0.003642349,0.0010098846,0.0017923018,0.002200646,0.0016672403,0.00081314927],"category_scores_gemma":[0.051598813,0.0005088734,0.00095012865,0.0027940592,0.0018239921,0.004386987,0.002857303,0.0022259993,0.0004873303],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008631826,0.0005521677,0.5187277,0.0014780269,0.0004068712,0.00081254315,0.0031032385,0.22251256,0.008665747,0.00905852,0.015963733,0.21785575],"study_design_scores_gemma":[0.00010119713,0.00064212683,0.123995304,0.00025211385,0.00015413997,0.0021525186,0.0024291503,0.8063802,0.0075435727,0.04429926,0.011927992,0.00012240966],"about_ca_topic_score_codex":0.005835345,"about_ca_topic_score_gemma":0.009242274,"teacher_disagreement_score":0.0064547565,"about_ca_system_score_codex":0.0012227026,"about_ca_system_score_gemma":0.0014266805,"threshold_uncertainty_score":0.034136415},"labels":[],"label_agreement":null},{"id":"W7126425372","doi":"10.21428/594757db.88028518","title":"C++ Source Code Verification Pipeline","year":2025,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Bespoke; Source code; Correctness; Classifier (UML); Pipeline (software); Artificial neural network; Software; Deep learning; Natural language","score_opus":0.01443259837562605,"score_gpt":0.27954924439643625,"score_spread":0.2651166460208102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126425372","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021624895,0.000490263,0.4352247,0.0006685465,0.0002883343,0.0007352447,0.016245378,0.5080336,0.016689023],"genre_scores_gemma":[0.3289767,0.000565762,0.47162086,0.0012146886,0.00014613343,0.0009255805,0.09392109,0.060986985,0.041642286],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984426,0.00011780873,0.00009987623,0.0004835965,0.00070143794,0.0001548066],"domain_scores_gemma":[0.99642974,0.00078173773,0.00023385423,0.0007865852,0.0016319242,0.00013607749],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011548041,0.0016199071,0.0008003051,0.0022024056,0.0005579444,0.001499817,0.002136952,0.00071719853,0.040538568],"category_scores_gemma":[0.009157396,0.0007290137,0.0013607331,0.0010521766,0.0004974088,0.001966233,0.0019533872,0.0018205227,0.02625184],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007816142,0.00025407993,0.006700076,0.0011276787,0.00016930343,0.0007773848,0.00038567052,0.02976793,0.029299762,0.015049489,0.25258142,0.6631056],"study_design_scores_gemma":[0.00031879323,0.00041674823,0.009159341,0.0002940168,0.00008192028,0.0009751784,0.00022898921,0.5563845,0.14126396,0.034201562,0.25645596,0.00021907601],"about_ca_topic_score_codex":0.0071231145,"about_ca_topic_score_gemma":0.0067282836,"teacher_disagreement_score":0.040538568,"about_ca_system_score_codex":0.0012981595,"about_ca_system_score_gemma":0.0036412247,"threshold_uncertainty_score":0.13561505},"labels":[],"label_agreement":null},{"id":"W7127277507","doi":"10.1109/ccece64018.2025.11364364","title":"AI-Assisted System Design: Improving Sequence Diagrams Through Use Case Scenarios","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University","funders":"","keywords":"Workflow; Sequence diagram; Leverage (statistics); Pipeline (software); Pipeline transport; Software; Consistency (knowledge bases); Sequence (biology)","score_opus":0.0781655486739429,"score_gpt":0.3159941776659975,"score_spread":0.23782862899205462,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7127277507","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024985814,0.00013176845,0.9667585,0.00042768675,0.000025053989,0.0004815176,0.0002186195,0.0037390087,0.0032320088],"genre_scores_gemma":[0.10105272,0.00013091021,0.8968185,0.00006681045,0.0000073905285,0.0002786419,0.000549107,0.0003702706,0.0007257363],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9841293,0.010947446,0.0007795515,0.0010940421,0.002688215,0.00036142147],"domain_scores_gemma":[0.9271143,0.05297132,0.0028886104,0.0101892445,0.0060870326,0.0007495633],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013423896,0.0019169305,0.0005447875,0.0044232323,0.0008685961,0.0044053383,0.0026830188,0.0015036373,0.0039901477],"category_scores_gemma":[0.070325576,0.001124833,0.0014192666,0.001994408,0.0016786244,0.0074996096,0.0038160179,0.0021850865,0.0013797074],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064243557,0.0011440781,0.008161778,0.0014505744,0.0002188618,0.0011528698,0.021099158,0.1343816,0.038514104,0.07713655,0.006933922,0.7091641],"study_design_scores_gemma":[0.00022214862,0.0005448882,0.002695303,0.0006952179,0.00015007924,0.0010808627,0.004140978,0.72352177,0.06688888,0.10948666,0.09035174,0.00022153129],"about_ca_topic_score_codex":0.0028849822,"about_ca_topic_score_gemma":0.0038418446,"teacher_disagreement_score":0.013423896,"about_ca_system_score_codex":0.001449865,"about_ca_system_score_gemma":0.0026369684,"threshold_uncertainty_score":0.070993185},"labels":[],"label_agreement":null},{"id":"W7127468921","doi":"10.1109/ccece64018.2025.11364388","title":"The Medley Interlisp Project: Reviving a Historical Software System","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Blueprint; Software; Software system; USable; Software development","score_opus":0.020084876299143094,"score_gpt":0.27868711797489815,"score_spread":0.2586022416757551,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7127468921","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11266908,0.015187465,0.45785922,0.08085588,0.0064783194,0.00034507518,0.00095591956,0.018830093,0.30681896],"genre_scores_gemma":[0.4030501,0.013254799,0.31564343,0.0094954,0.0012448669,0.00044602103,0.0023037891,0.0083538685,0.24620782],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99626654,0.0013070658,0.00015845872,0.00025941036,0.0016771049,0.00033145028],"domain_scores_gemma":[0.9933333,0.0013796346,0.00061009725,0.0016976076,0.0018805611,0.0010988323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064777927,0.00041033342,0.00016311002,0.0015487522,0.003001696,0.00517747,0.0017255635,0.0013522131,0.008687669],"category_scores_gemma":[0.013865734,0.0002811902,0.00026686204,0.0012400564,0.0045480905,0.008490211,0.0064481846,0.0035960963,0.0022311066],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001500405,0.00007980553,0.0013803843,0.00031875062,0.000012998521,0.00033624357,0.005134841,0.0019424444,0.0021871056,0.37279707,0.18026252,0.43539786],"study_design_scores_gemma":[0.000015200207,0.00008327069,0.0007339872,0.00028111122,0.000007897361,0.00044500467,0.0010157079,0.0022939844,0.0031605433,0.030625733,0.96130973,0.000027841506],"about_ca_topic_score_codex":0.0039745965,"about_ca_topic_score_gemma":0.0039083487,"teacher_disagreement_score":0.008687669,"about_ca_system_score_codex":0.0034150486,"about_ca_system_score_gemma":0.0065307515,"threshold_uncertainty_score":0.034258246},"labels":[],"label_agreement":null},{"id":"W7128625888","doi":"10.1109/icces51350.2021.11391692","title":"Retraction Notice: An Adaptable and Extensible Code Smell Detection Approach","year":2021,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Code smell; Code refactoring; Mistake; Code (set theory); Product (mathematics); Source code","score_opus":0.038857986760423104,"score_gpt":0.27653813082010664,"score_spread":0.23768014405968355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7128625888","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008069238,0.001025707,0.008602796,0.25575125,0.71846235,0.00014997661,0.0011141155,0.0016574836,0.012429462],"genre_scores_gemma":[0.025479434,0.003956022,0.028256526,0.19765718,0.41779608,0.0007073347,0.0028777067,0.00276604,0.32050377],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99016106,0.001466419,0.0012323991,0.0010714531,0.0054875566,0.00058107165],"domain_scores_gemma":[0.9387645,0.017275933,0.0026092087,0.004735116,0.033060037,0.003555278],"candidate_categories":["metaresearch","research_integrity"],"consensus_categories":[],"category_scores_codex":[0.008276798,0.0014576287,0.0007615939,0.0023624506,0.0034668518,0.005158221,0.0033443752,0.009818153,0.03624978],"category_scores_gemma":[0.11422068,0.0006200257,0.0015149169,0.0013076987,0.002802744,0.0038338818,0.00376229,0.015580501,0.029887745],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000166507,0.000010349348,0.000107383006,0.00007722384,0.0000064740725,0.00015563802,0.00008886002,0.000053485266,0.00019456571,0.0028489183,0.98354036,0.012900066],"study_design_scores_gemma":[0.000012011133,0.000028871274,0.00050652697,0.00015746136,0.000015311061,0.00029955554,0.00010231705,0.00056769344,0.00035645656,0.0018914059,0.99603397,0.000028470084],"about_ca_topic_score_codex":0.008609392,"about_ca_topic_score_gemma":0.008404668,"teacher_disagreement_score":0.9917232,"about_ca_system_score_codex":0.0037959374,"about_ca_system_score_gemma":0.00597895,"threshold_uncertainty_score":0.12126762},"labels":[{"model":"gemma","categories":["research_integrity"],"domain":null,"study_design":"not_applicable","genre":"editorial","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":["research_integrity"],"domain":null,"study_design":"not_applicable","genre":"editorial","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"agree"},{"id":"W7128994509","doi":"10.1109/icsess67729.2025.11380541","title":"Understanding the Amount of Changes Required for Merge Request Acceptance: An Empirical Study","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Telus (Canada); École de Technologie Supérieure","funders":"","keywords":"Merge (version control); Source lines of code; Code review; Empirical research; Code (set theory); Source code; Software; Software inspection","score_opus":0.19299391524720844,"score_gpt":0.40722450241815694,"score_spread":0.2142305871709485,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7128994509","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99559903,0.00034588977,0.0015992745,0.00026740384,0.00001194358,0.000102410384,0.0008546069,0.000106802996,0.0011125164],"genre_scores_gemma":[0.99363786,0.0002127732,0.0034824472,0.00009518048,0.000029437322,0.00013923421,0.0017380823,0.000055828987,0.00060911704],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.983273,0.007034453,0.0024109515,0.0022712082,0.004377672,0.0006326526],"domain_scores_gemma":[0.49771369,0.36789906,0.07912697,0.011452204,0.03762863,0.0061794426],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.023810947,0.00055665284,0.00038089626,0.004499834,0.00070932484,0.002323451,0.0013812325,0.0009114328,0.0020660008],"category_scores_gemma":[0.19467592,0.00040049863,0.00045374746,0.00407002,0.0009487478,0.0040754215,0.0013895279,0.001643199,0.0011034533],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047166165,0.0010256756,0.9218643,0.0005842346,0.00014949965,0.00030634613,0.0049320916,0.0031451893,0.0012691956,0.00040168772,0.0037623304,0.06208772],"study_design_scores_gemma":[0.000056179902,0.00080526463,0.9633315,0.00019910553,0.00009401815,0.0004469488,0.004256439,0.021482155,0.0018412496,0.0004316495,0.006975754,0.000079707446],"about_ca_topic_score_codex":0.0045771175,"about_ca_topic_score_gemma":0.0069498513,"teacher_disagreement_score":0.9761891,"about_ca_system_score_codex":0.0015723394,"about_ca_system_score_gemma":0.0016044276,"threshold_uncertainty_score":0.12592584},"labels":[],"label_agreement":null},{"id":"W7130731024","doi":"10.1109/swc65939.2025.00299","title":"Time Travel: LLM-Assisted Semantic Behavior Localization with Git Bisect","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University; University of Waterloo","funders":"","keywords":"Workflow; Annotation; Software; Semantics (computer science); Process (computing)","score_opus":0.014813862988059827,"score_gpt":0.26548442938142713,"score_spread":0.2506705663933673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7130731024","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010380685,0.000268471,0.7488756,0.00041582293,0.00011557217,0.0001329251,0.0026739386,0.23548976,0.001647249],"genre_scores_gemma":[0.19780524,0.00020370472,0.76165235,0.0006487424,0.000062290404,0.00038865508,0.012299316,0.022848424,0.0040912824],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969332,0.0007929819,0.00023708395,0.0008573937,0.0009962292,0.00018307398],"domain_scores_gemma":[0.9922225,0.0029389102,0.00069510913,0.002731465,0.0011573054,0.00025477356],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028387387,0.0024142552,0.0009598438,0.0024620255,0.0008541706,0.002926576,0.003717344,0.0016831069,0.006868028],"category_scores_gemma":[0.019463923,0.0011819567,0.002247353,0.0013978673,0.0015194749,0.0057258443,0.0056214062,0.003563728,0.0049681556],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014439666,0.0005526346,0.023011975,0.0020579388,0.00046812207,0.001016821,0.004116731,0.1312419,0.040255643,0.060288455,0.11711606,0.6184298],"study_design_scores_gemma":[0.00009062823,0.00012527764,0.0011782076,0.000118710974,0.00009219511,0.00021371279,0.00036483855,0.878813,0.028285304,0.053401913,0.03721279,0.000103297156],"about_ca_topic_score_codex":0.010126521,"about_ca_topic_score_gemma":0.021333601,"teacher_disagreement_score":0.010126521,"about_ca_system_score_codex":0.001619972,"about_ca_system_score_gemma":0.003860368,"threshold_uncertainty_score":0.022975862},"labels":[],"label_agreement":null},{"id":"W7131392892","doi":"10.1109/icdm65498.2025.00138","title":"SBAN: A Framework &amp; Multi-Dimensional Dataset for Large Language Model Pre-Training and Software Code Mining","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Natural language; Malware; Code review; Software; Software construction; Software development; Source code; Bridging (networking); Software system; Software framework","score_opus":0.05074200444759836,"score_gpt":0.36094767752618245,"score_spread":0.3102056730785841,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7131392892","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038916055,0.0021095304,0.3913822,0.0032505821,0.00075678475,0.0017212009,0.43891153,0.11658603,0.006366088],"genre_scores_gemma":[0.051826045,0.00044072702,0.31945038,0.00074536505,0.00010057227,0.0025836932,0.6210698,0.001678765,0.0021045771],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99733514,0.00074369396,0.00029941215,0.0007469302,0.00067318603,0.00020162654],"domain_scores_gemma":[0.9954159,0.0020516643,0.0003416275,0.001267739,0.0006641623,0.00025892202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031280303,0.0027397338,0.0009801044,0.004912974,0.001339357,0.001955855,0.0045660944,0.0024295119,0.0046817493],"category_scores_gemma":[0.015157667,0.0008831607,0.0022919725,0.003527551,0.0011005171,0.0033265515,0.0045669805,0.0037225254,0.0056344527],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007664424,0.0010795641,0.01738783,0.0020872944,0.00053836725,0.0007272964,0.00070702995,0.052611396,0.009869468,0.020381562,0.59432125,0.29952246],"study_design_scores_gemma":[0.00045843047,0.00041045874,0.009889265,0.00034343978,0.00015579184,0.00079642696,0.000585538,0.625846,0.015551076,0.05175665,0.29395196,0.0002549097],"about_ca_topic_score_codex":0.0146690225,"about_ca_topic_score_gemma":0.032723915,"teacher_disagreement_score":0.0146690225,"about_ca_system_score_codex":0.0018392745,"about_ca_system_score_gemma":0.002891321,"threshold_uncertainty_score":0.029167295},"labels":[],"label_agreement":null},{"id":"W7133323927","doi":"10.5281/zenodo.18845271","title":"OSSVul - ReplicationPackage","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Scripting language; Software; Vulnerability (computing); Artifact (error); Component (thermodynamics); Identification (biology); Data collection; Timestamp","score_opus":0.021157852431734268,"score_gpt":0.25365114832424446,"score_spread":0.2324932958925102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7133323927","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0332162,0.0006891321,0.09142034,0.0007976625,0.00040781117,0.002146033,0.5164535,0.33443996,0.020429319],"genre_scores_gemma":[0.067750834,0.00039006877,0.087038,0.00028928378,0.000083237246,0.0030298876,0.8162466,0.020686788,0.0044853375],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99687237,0.0008042253,0.00042944116,0.0007663553,0.000930067,0.00019755602],"domain_scores_gemma":[0.99005353,0.0026915174,0.00046107845,0.004686253,0.0018350896,0.00027251808],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0053053233,0.0020758347,0.0010558782,0.0035473695,0.0007576448,0.0028145062,0.0031940069,0.00087174523,0.020869128],"category_scores_gemma":[0.022722533,0.0009504547,0.0020562068,0.00300499,0.000642816,0.003092003,0.0029719637,0.0017127906,0.014139712],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011438663,0.0005673841,0.019199828,0.002495756,0.0004707132,0.0002501953,0.0006509668,0.052240428,0.0037991055,0.01627783,0.8017792,0.10112464],"study_design_scores_gemma":[0.0008812541,0.0006952852,0.019056078,0.00063588924,0.00018928818,0.00037035844,0.0007088819,0.24418452,0.0138465185,0.0324896,0.6866734,0.00026890164],"about_ca_topic_score_codex":0.0118781235,"about_ca_topic_score_gemma":0.009895256,"teacher_disagreement_score":0.99469465,"about_ca_system_score_codex":0.0012845901,"about_ca_system_score_gemma":0.0022471072,"threshold_uncertainty_score":0.069814205},"labels":[],"label_agreement":null},{"id":"W7134898377","doi":"10.1109/icdmw69685.2025.00152","title":"SBAN: A Framework &amp; Multi-Dimensional Dataset for Large Language Model Pre-Training and Software Code Mining","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Software; Code (set theory); Source code; Software system; Modeling language; Key (lock); Data modeling","score_opus":0.05074200444759836,"score_gpt":0.36094767752618245,"score_spread":0.3102056730785841,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7134898377","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038916055,0.0021095304,0.3913822,0.0032505821,0.00075678475,0.0017212009,0.43891153,0.11658603,0.006366088],"genre_scores_gemma":[0.051826045,0.00044072702,0.31945038,0.00074536505,0.00010057227,0.0025836932,0.6210698,0.001678765,0.0021045771],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99733514,0.00074369396,0.00029941215,0.0007469302,0.00067318603,0.00020162654],"domain_scores_gemma":[0.9954159,0.0020516643,0.0003416275,0.001267739,0.0006641623,0.00025892202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031280303,0.0027397338,0.0009801044,0.004912974,0.001339357,0.001955855,0.0045660944,0.0024295119,0.0046817493],"category_scores_gemma":[0.015157667,0.0008831607,0.0022919725,0.003527551,0.0011005171,0.0033265515,0.0045669805,0.0037225254,0.0056344527],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007664424,0.0010795641,0.01738783,0.0020872944,0.00053836725,0.0007272964,0.00070702995,0.052611396,0.009869468,0.020381562,0.59432125,0.29952246],"study_design_scores_gemma":[0.00045843047,0.00041045874,0.009889265,0.00034343978,0.00015579184,0.00079642696,0.000585538,0.625846,0.015551076,0.05175665,0.29395196,0.0002549097],"about_ca_topic_score_codex":0.0146690225,"about_ca_topic_score_gemma":0.032723915,"teacher_disagreement_score":0.0146690225,"about_ca_system_score_codex":0.0018392745,"about_ca_system_score_gemma":0.002891321,"threshold_uncertainty_score":0.029167295},"labels":[],"label_agreement":null},{"id":"W7134957359","doi":"10.1109/icdmw69685.2025.00164","title":"Security Charter Effectiveness in Large Language Model Code Generation: A Multi-Phase Experimental Analysis Revealing Task-Dependent Responsiveness and Architectural Differences","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute","funders":"","keywords":"Charter; Code (set theory); Key (lock); Modeling language; Semantics (computer science); Code-switching","score_opus":0.026207704539285115,"score_gpt":0.3530747111129729,"score_spread":0.3268670065736878,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7134957359","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9665195,0.00046287963,0.025504446,0.00031032824,0.00014407965,0.00073859503,0.0003971262,0.0024073452,0.0035157437],"genre_scores_gemma":[0.9703159,0.00014494163,0.023990193,0.0003095115,0.000047515394,0.0012156339,0.00091495045,0.0006789628,0.0023823085],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9868117,0.008537625,0.0010139181,0.0019162826,0.0013475778,0.00037293395],"domain_scores_gemma":[0.8881164,0.09551155,0.004373876,0.008387701,0.002287633,0.0013227873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018081833,0.0010992995,0.00082004967,0.00060639,0.00039848572,0.0020157392,0.0012941813,0.0010088999,0.0024324905],"category_scores_gemma":[0.1148859,0.0005552809,0.0006397369,0.00039942915,0.0010817054,0.00303472,0.0020523064,0.0022302798,0.0011449013],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.043922026,0.01061632,0.04998676,0.0022829687,0.0012798253,0.00034376612,0.008778602,0.0934465,0.12603761,0.0049767517,0.013611488,0.64471745],"study_design_scores_gemma":[0.008824706,0.09083862,0.15427396,0.0006411987,0.001411685,0.0006979792,0.0036385758,0.51006854,0.17323561,0.021277351,0.034105502,0.0009862178],"about_ca_topic_score_codex":0.0010916294,"about_ca_topic_score_gemma":0.0009932792,"teacher_disagreement_score":0.018081833,"about_ca_system_score_codex":0.0007092148,"about_ca_system_score_gemma":0.0010536822,"threshold_uncertainty_score":0.09562701},"labels":[],"label_agreement":null},{"id":"W7135003010","doi":"","title":"Exploring Cross-Language Software Similarity Analysis Using Source Code Context","year":2025,"lang":"en","type":"article","venue":"University Library (University of Saskatchewan)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Software development; Source code; Software construction; Static program analysis; Code review; Software maintenance; Documentation; Software quality; Code reuse; Software portability","score_opus":0.027950171179192157,"score_gpt":0.22329555196826295,"score_spread":0.1953453807890708,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7135003010","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58245516,0.03237886,0.3627253,0.0021689956,0.00021892249,0.0005973407,0.00557787,0.002530978,0.011346545],"genre_scores_gemma":[0.8606523,0.0039109346,0.12719981,0.00039381607,0.00008507007,0.0004499387,0.0060423985,0.00033316255,0.00093264855],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9900568,0.0032290616,0.0011230442,0.0026230374,0.002663901,0.0003041799],"domain_scores_gemma":[0.94362456,0.036320277,0.0080217905,0.004340455,0.0070983777,0.0005945201],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007051645,0.0007986196,0.00087687315,0.022825703,0.0010855534,0.00364541,0.0016209933,0.0011637034,0.0008975087],"category_scores_gemma":[0.059631206,0.00043773756,0.0013116796,0.012944662,0.0014927044,0.007996267,0.00411591,0.001182326,0.0004035604],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025711083,0.00027658223,0.24074563,0.0066573704,0.0010069216,0.0012899817,0.011434344,0.0068712626,0.010612867,0.03454724,0.007435887,0.6788647],"study_design_scores_gemma":[0.0001349339,0.0007701596,0.3458757,0.006410955,0.002456357,0.003971956,0.027782628,0.24564502,0.028157698,0.19031337,0.14798872,0.000492468],"about_ca_topic_score_codex":0.0048427293,"about_ca_topic_score_gemma":0.00785405,"teacher_disagreement_score":0.022825703,"about_ca_system_score_codex":0.0012454784,"about_ca_system_score_gemma":0.0030697072,"threshold_uncertainty_score":0.037293136},"labels":[],"label_agreement":null},{"id":"W7135874681","doi":"","title":"Can Genetic Programming improve Software Effort Estimation? A Comparative Evaluation","year":2005,"lang":"en","type":"book-chapter","venue":"Explore Bristol Research","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Software; Genetic programming; Software development; Software metric; Variety (cybernetics); Set (abstract data type); Software sizing; Software construction","score_opus":0.15587769726479406,"score_gpt":0.3972212044288398,"score_spread":0.24134350716404573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7135874681","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40352002,0.07976056,0.3467785,0.007697376,0.00065150595,0.00031536183,0.00080571714,0.0018590556,0.15861192],"genre_scores_gemma":[0.70631856,0.031484645,0.25329864,0.000580131,0.00027376102,0.00014935901,0.00087211566,0.00040559235,0.006617282],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9922307,0.0046104956,0.00021064578,0.00048152413,0.002310517,0.0001561312],"domain_scores_gemma":[0.9650235,0.031328086,0.0004789045,0.0011039516,0.0019100967,0.00015534688],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009050007,0.0009731028,0.00085708086,0.0034201092,0.0002802949,0.0015195938,0.0015957686,0.0014245008,0.0023914455],"category_scores_gemma":[0.036217857,0.00019014518,0.0007218228,0.006640056,0.0007004908,0.0024330772,0.00088023546,0.0007889494,0.0004492983],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043402874,0.0002473302,0.005243651,0.00085945975,0.0002728677,0.00008734935,0.00025757862,0.095127635,0.00051788107,0.01962547,0.0028011394,0.8745256],"study_design_scores_gemma":[0.00042043257,0.0041724173,0.04175006,0.0024837658,0.0012115949,0.0005579718,0.0013086005,0.79487795,0.011525435,0.081807464,0.059711374,0.00017289126],"about_ca_topic_score_codex":0.0029645409,"about_ca_topic_score_gemma":0.003137063,"teacher_disagreement_score":0.009050007,"about_ca_system_score_codex":0.0017557658,"about_ca_system_score_gemma":0.0008200899,"threshold_uncertainty_score":0.047861636},"labels":[],"label_agreement":null},{"id":"W7139036677","doi":"10.1109/adacis65663.2025.11436700","title":"A Survey on Co-occurrences of Code and Test Smells","year":2025,"lang":"","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Code smell; Code (set theory); Code review; Software; Code refactoring; Test (biology)","score_opus":0.04282926802888274,"score_gpt":0.3387565119823106,"score_spread":0.29592724395342784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7139036677","genre_codex":"empirical","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5349914,0.40479112,0.021402152,0.004869934,0.00050880836,0.00035818195,0.0029576414,0.0006441228,0.029476687],"genre_scores_gemma":[0.6904487,0.28064123,0.017285755,0.0017436643,0.00036308865,0.0004331797,0.0043747844,0.00036782946,0.0043417206],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.974161,0.007817562,0.0050355573,0.0020156293,0.010273344,0.00069681945],"domain_scores_gemma":[0.7703171,0.16828844,0.029876309,0.004769624,0.024843972,0.001904611],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01251812,0.0006008029,0.0008948117,0.027923081,0.0009611657,0.0031298331,0.00092035346,0.001148384,0.0022276074],"category_scores_gemma":[0.06634096,0.0005978796,0.0011019255,0.02882861,0.0014975208,0.007086852,0.00279475,0.0009258854,0.000684504],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021679046,0.00013714415,0.31069982,0.0136267245,0.00030610675,0.0013277796,0.031894166,0.0005047453,0.005027188,0.004179198,0.007943671,0.6241366],"study_design_scores_gemma":[0.000022213015,0.0008346859,0.63956755,0.02123198,0.00044339264,0.010431824,0.053298317,0.0010437391,0.0060286825,0.0039267447,0.26293194,0.00023902867],"about_ca_topic_score_codex":0.0035218769,"about_ca_topic_score_gemma":0.004038206,"teacher_disagreement_score":0.027923081,"about_ca_system_score_codex":0.0013725699,"about_ca_system_score_gemma":0.0020263782,"threshold_uncertainty_score":0.06620294},"labels":[],"label_agreement":null},{"id":"W7151437856","doi":"10.5281/zenodo.19449192","title":"Advancements in Early-Stage Software Vulnerability Detection","year":2017,"lang":"","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cybernet Systems Corporation (Canada)","funders":"","keywords":"Software; Fault (geology); Construct (python library); Code (set theory); Fault detection and isolation; Source code; Data modeling; Fault model","score_opus":0.04263471353159285,"score_gpt":0.2828024519658255,"score_spread":0.24016773843423264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7151437856","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36411536,0.012225418,0.5925275,0.0031626576,0.0004089494,0.00032017802,0.010355373,0.007048916,0.009835623],"genre_scores_gemma":[0.7125065,0.0044002747,0.25934678,0.0004442196,0.00027933632,0.0002084282,0.018784426,0.0003510914,0.0036789416],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9930768,0.0020180908,0.0004716178,0.0019251816,0.0023142775,0.00019404574],"domain_scores_gemma":[0.9694056,0.014783578,0.00243559,0.0070505706,0.005670617,0.00065398397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059187235,0.001587584,0.0012902182,0.0055689425,0.00044914478,0.0021036032,0.0024124403,0.0013166894,0.0012199443],"category_scores_gemma":[0.023116862,0.0005750648,0.0015447219,0.0039798478,0.00062355114,0.0041295732,0.0018163016,0.0031656977,0.0018941893],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002502205,0.0011426692,0.118832655,0.0009522975,0.00042974117,0.00010375887,0.00033964458,0.04192959,0.007928562,0.0034364907,0.011122561,0.8135318],"study_design_scores_gemma":[0.00006363124,0.0011606376,0.13102753,0.00062574906,0.00026332017,0.000777171,0.00047287965,0.75213635,0.02953519,0.024933703,0.05882111,0.0001828079],"about_ca_topic_score_codex":0.00401379,"about_ca_topic_score_gemma":0.004135906,"teacher_disagreement_score":0.0059187235,"about_ca_system_score_codex":0.00074679113,"about_ca_system_score_gemma":0.0011304569,"threshold_uncertainty_score":0.031301618},"labels":[],"label_agreement":null},{"id":"W7151460534","doi":"10.5281/zenodo.19449191","title":"Advancements in Early-Stage Software Vulnerability Detection","year":2017,"lang":"","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cybernet Systems Corporation (Canada)","funders":"","keywords":"Software; Fault (geology); Construct (python library); Code (set theory); Fault detection and isolation; Source code; Data modeling; Fault model","score_opus":0.04263471353159285,"score_gpt":0.2828024519658255,"score_spread":0.24016773843423264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7151460534","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36411536,0.012225418,0.5925275,0.0031626576,0.0004089494,0.00032017802,0.010355373,0.007048916,0.009835623],"genre_scores_gemma":[0.7125065,0.0044002747,0.25934678,0.0004442196,0.00027933632,0.0002084282,0.018784426,0.0003510914,0.0036789416],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9930768,0.0020180908,0.0004716178,0.0019251816,0.0023142775,0.00019404574],"domain_scores_gemma":[0.9694056,0.014783578,0.00243559,0.0070505706,0.005670617,0.00065398397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059187235,0.001587584,0.0012902182,0.0055689425,0.00044914478,0.0021036032,0.0024124403,0.0013166894,0.0012199443],"category_scores_gemma":[0.023116862,0.0005750648,0.0015447219,0.0039798478,0.00062355114,0.0041295732,0.0018163016,0.0031656977,0.0018941893],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002502205,0.0011426692,0.118832655,0.0009522975,0.00042974117,0.00010375887,0.00033964458,0.04192959,0.007928562,0.0034364907,0.011122561,0.8135318],"study_design_scores_gemma":[0.00006363124,0.0011606376,0.13102753,0.00062574906,0.00026332017,0.000777171,0.00047287965,0.75213635,0.02953519,0.024933703,0.05882111,0.0001828079],"about_ca_topic_score_codex":0.00401379,"about_ca_topic_score_gemma":0.004135906,"teacher_disagreement_score":0.0059187235,"about_ca_system_score_codex":0.00074679113,"about_ca_system_score_gemma":0.0011304569,"threshold_uncertainty_score":0.031301618},"labels":[],"label_agreement":null},{"id":"W78107666","doi":"10.1007/978-3-642-01680-6_2","title":"Decision Processes for Trustworthy Software","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Transparency (behavior); Computer science; Trustworthiness; Software; Decision support system; Key (lock); Management science; Empirical research; Software engineering; Knowledge management; Process management; Data mining; Engineering; Computer security","score_opus":0.018302708062149838,"score_gpt":0.2726176058727606,"score_spread":0.2543148978106108,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W78107666","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054414338,0.0009641626,0.97787946,0.00084349373,0.00011289561,0.000102719125,0.000040225925,0.00016206302,0.014453567],"genre_scores_gemma":[0.3313727,0.0020472615,0.64536935,0.00030637253,0.00030615125,0.00055897265,0.00021436145,0.0001993124,0.019625567],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99292314,0.00279968,0.00057622895,0.00081331044,0.0024463888,0.00044123447],"domain_scores_gemma":[0.9889818,0.008201616,0.0005721044,0.0012314867,0.0007826872,0.00023021191],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065661683,0.0012483716,0.0011892477,0.0011649445,0.0012216845,0.0043153805,0.001829258,0.0022237326,0.009119672],"category_scores_gemma":[0.020697562,0.00095059595,0.0016851905,0.0015117811,0.0037303467,0.005561854,0.0030301085,0.0047106766,0.001723069],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046850804,0.000035115678,0.00013478247,0.0001647639,0.00003345665,0.000062891755,0.00037993403,0.020283528,0.0006165297,0.92030513,0.0014804689,0.05645661],"study_design_scores_gemma":[0.000015727868,0.000024692494,0.000046351885,0.000051101375,0.0000129471055,0.00002674972,0.000046159377,0.068908155,0.0005150802,0.9262377,0.0041037463,0.0000115744415],"about_ca_topic_score_codex":0.0015010507,"about_ca_topic_score_gemma":0.0011624375,"teacher_disagreement_score":0.009119672,"about_ca_system_score_codex":0.0022161207,"about_ca_system_score_gemma":0.0024554536,"threshold_uncertainty_score":0.034725606},"labels":[],"label_agreement":null},{"id":"W79425221","doi":"","title":"A fuzzy-based multimodel system for reasoning about the number of software defects: Research Articles","year":2005,"lang":"en","type":"article","venue":"International Journal of Intelligent Systems","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Fuzzy logic; Software; Artificial intelligence; Function (biology); Machine learning; Data mining; Software engineering","score_opus":0.0621005207150327,"score_gpt":0.3703374831358635,"score_spread":0.3082369624208308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W79425221","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011987918,0.0004194348,0.9839107,0.00035387604,0.000038473132,0.0001036559,0.0002597297,0.0012017628,0.0017244412],"genre_scores_gemma":[0.25632828,0.0006363201,0.73995185,0.0001689314,0.000075149386,0.00027145707,0.00060926843,0.00007855485,0.0018801293],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99901617,0.00023949267,0.00013213675,0.00023488668,0.0003256159,0.000051649844],"domain_scores_gemma":[0.9980386,0.0010983349,0.00015821568,0.00024580007,0.0003872124,0.00007179975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002192059,0.0008134181,0.0010079454,0.0027347903,0.00086525193,0.0030730462,0.0018454952,0.0020599235,0.004983946],"category_scores_gemma":[0.007960217,0.00042311064,0.0013710617,0.0016230451,0.0006374752,0.003234681,0.0010363187,0.001037643,0.0010912752],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037555405,0.00038718528,0.005473255,0.0005830527,0.00035932058,0.0007213097,0.00092903554,0.25307813,0.011945288,0.06437793,0.0051386007,0.65663135],"study_design_scores_gemma":[0.000023616314,0.00005433019,0.00049539463,0.00007505192,0.00007968264,0.00017558149,0.00006950461,0.9531047,0.002130882,0.039144758,0.004613455,0.000033085496],"about_ca_topic_score_codex":0.00661231,"about_ca_topic_score_gemma":0.0068255737,"teacher_disagreement_score":0.00661231,"about_ca_system_score_codex":0.0013049627,"about_ca_system_score_gemma":0.0012967952,"threshold_uncertainty_score":0.01667291},"labels":[],"label_agreement":null},{"id":"W79589210","doi":"","title":"ON THE DISTRIBUTION OF SOURCE CODE FILE SIZES","year":2011,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Victoria","funders":"","keywords":"Computer science; Source code; Log-normal distribution; Software; Code (set theory); Software quality; Pareto distribution; Statistics; Operating system; Mathematics; Software development; Programming language; Set (abstract data type)","score_opus":0.030093604044607904,"score_gpt":0.2369343151585387,"score_spread":0.2068407111139308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W79589210","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9494039,0.00046125535,0.044424243,0.00058506464,0.000024792811,0.000055282824,0.00124891,0.0004271887,0.0033693954],"genre_scores_gemma":[0.99440914,0.00017102069,0.003486502,0.000038529724,0.00003539502,0.00004686454,0.0012623155,0.000077730256,0.00047242374],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99627954,0.0012313682,0.0001705107,0.0009402855,0.0010803294,0.00029793114],"domain_scores_gemma":[0.87731713,0.09060915,0.009505787,0.0069524334,0.014083991,0.0015315863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072556334,0.0005762368,0.0004156758,0.0045724493,0.00041007614,0.0017174075,0.0013144354,0.0013380534,0.0022968387],"category_scores_gemma":[0.08253697,0.00047740334,0.00036772713,0.0027694507,0.0018659267,0.0032366554,0.0010999228,0.0010356945,0.00075873546],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010155446,0.00030196112,0.6052814,0.00031601553,0.00020603863,0.00085910945,0.001513054,0.25476578,0.010380688,0.039220892,0.0056350213,0.080504484],"study_design_scores_gemma":[0.000063183805,0.0002154547,0.2456012,0.00013551881,0.000056813304,0.0013027517,0.0005727441,0.71908337,0.007347787,0.022870198,0.0026246193,0.00012637333],"about_ca_topic_score_codex":0.0032421614,"about_ca_topic_score_gemma":0.0017655216,"teacher_disagreement_score":0.0072556334,"about_ca_system_score_codex":0.0017982149,"about_ca_system_score_gemma":0.0005965661,"threshold_uncertainty_score":0.03837192},"labels":[],"label_agreement":null},{"id":"W809289544","doi":"10.71781/15153","title":"Rule-based quality heuristics formalization and identification","year":2007,"lang":"en","type":"dissertation","venue":"Papyrus : Institutional Repository (Université de Montréal)","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Université de Montréal","keywords":"Heuristics; Computer science; Software quality; Software; Software engineering; Data mining; Quality (philosophy); Programming language; Software development","score_opus":0.009245377595690055,"score_gpt":0.2294565300528428,"score_spread":0.22021115245715275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W809289544","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007883841,0.00039844462,0.9835926,0.00052397413,0.00007848085,0.00039954085,0.00039686493,0.0020315473,0.0046945433],"genre_scores_gemma":[0.14209667,0.00032858626,0.8515743,0.00013865202,0.000064330474,0.00026646015,0.0010016953,0.00046457915,0.004064753],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9883052,0.0034212105,0.0013382551,0.0014910115,0.0047456855,0.00069855514],"domain_scores_gemma":[0.9691486,0.017839836,0.0018262629,0.004475729,0.006250726,0.0004589856],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009889118,0.0009119131,0.0012558701,0.003949586,0.0011123992,0.006147612,0.004175807,0.0015562789,0.005771427],"category_scores_gemma":[0.03709953,0.0012537935,0.0026370324,0.0025198057,0.002819155,0.0052318014,0.003227745,0.0028433898,0.0014267267],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003205604,0.00046802466,0.0047296374,0.0011285545,0.00029937242,0.0009282846,0.0026415733,0.10125974,0.010756656,0.47226155,0.014428355,0.39077774],"study_design_scores_gemma":[0.00017413478,0.00016489769,0.0020259032,0.0005690715,0.00034296254,0.0006648173,0.0009203441,0.6439684,0.027774489,0.26820654,0.055031173,0.00015732502],"about_ca_topic_score_codex":0.017486587,"about_ca_topic_score_gemma":0.022803683,"teacher_disagreement_score":0.017486587,"about_ca_system_score_codex":0.003684494,"about_ca_system_score_gemma":0.005917489,"threshold_uncertainty_score":0.05229926},"labels":[],"label_agreement":null},{"id":"W814172419","doi":"10.1109/saner.2016.105","title":"Analyzing the State of Static Analysis: A Large-Scale Evaluation in Open Source Software","year":2016,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":206,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Java; Python (programming language); Computer science; Open source software; Software; Software engineering; JavaScript; Open source; Static analysis; Population; Programming language","score_opus":0.026352838873044706,"score_gpt":0.31568346909399925,"score_spread":0.28933063022095457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W814172419","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9919837,0.00045047523,0.0051012025,0.00016429016,0.000017965676,0.00016160031,0.0004531488,0.000531563,0.0011360416],"genre_scores_gemma":[0.9885465,0.00021056732,0.008680785,0.00007668323,0.00001784792,0.00017975685,0.0016915688,0.00022744056,0.00036879416],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9716178,0.013762089,0.0017364315,0.0034762793,0.008642665,0.0007646695],"domain_scores_gemma":[0.756635,0.17709951,0.017491562,0.02026133,0.025314908,0.0031976933],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028521102,0.0007782911,0.00061885995,0.0059962925,0.0010994598,0.0017591277,0.0016860323,0.00091445114,0.0007381137],"category_scores_gemma":[0.1397941,0.00046795278,0.0008979283,0.00476601,0.0026153552,0.0043976465,0.0035596953,0.0017456248,0.00037802928],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019081498,0.0033704396,0.6419811,0.0019651935,0.0008463771,0.0004271337,0.021301556,0.010906737,0.00885995,0.0027637554,0.0068961927,0.2987735],"study_design_scores_gemma":[0.0003177937,0.003105235,0.8750153,0.00067103305,0.0005109677,0.0005368605,0.009835467,0.08266614,0.011603332,0.004469115,0.011067212,0.00020162914],"about_ca_topic_score_codex":0.0055456264,"about_ca_topic_score_gemma":0.006617051,"teacher_disagreement_score":0.9714789,"about_ca_system_score_codex":0.001975227,"about_ca_system_score_gemma":0.0015937088,"threshold_uncertainty_score":0.15083581},"labels":[],"label_agreement":null},{"id":"W8239271","doi":"10.1007/978-1-4614-6596-6_15","title":"Facilitating Crowd Sourced Software Engineering via Stack Overflow","year":2013,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Documentation; Computer science; World Wide Web; Source code; The Internet; Context (archaeology); Software; Code (set theory); Software engineering; Set (abstract data type); Programming language","score_opus":0.015922404427501972,"score_gpt":0.2201680234086281,"score_spread":0.20424561898112611,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W8239271","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018411132,0.0030206048,0.71484935,0.0013196435,0.0006009731,0.00018391693,0.00019208656,0.012426038,0.2489962],"genre_scores_gemma":[0.16344847,0.005626562,0.4921345,0.0007187747,0.0005603779,0.00024221788,0.001159588,0.00837087,0.32773867],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989213,0.00017861337,0.000028298902,0.00011332804,0.00065850624,0.00009995176],"domain_scores_gemma":[0.9982285,0.0008830871,0.00006180232,0.00045779103,0.00024404767,0.0001248558],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012697816,0.0010494835,0.00053611235,0.0015377766,0.0009139473,0.00294558,0.0017128828,0.001044371,0.025982905],"category_scores_gemma":[0.0051331837,0.00051813776,0.0006482756,0.0014225724,0.00080981763,0.0076006847,0.0059453286,0.0021691425,0.010097714],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044760785,0.00014887443,0.000358067,0.00029132023,0.000017409293,0.00026926023,0.0018186378,0.0039969534,0.009762576,0.087350205,0.046246108,0.8496958],"study_design_scores_gemma":[0.00003240183,0.00008566562,0.0011327728,0.0005787044,0.000052836633,0.0014169655,0.0010496906,0.037353992,0.026047077,0.26253408,0.66962767,0.000087977576],"about_ca_topic_score_codex":0.000739171,"about_ca_topic_score_gemma":0.0018911689,"teacher_disagreement_score":0.025982905,"about_ca_system_score_codex":0.0007769109,"about_ca_system_score_gemma":0.0012824574,"threshold_uncertainty_score":0.08692151},"labels":[],"label_agreement":null},{"id":"W833290532","doi":"10.71781/10471","title":"Modelling software quality : a multidimensional approach","year":2010,"lang":"en","type":"dissertation","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Quality (philosophy); Software quality; Component (thermodynamics); Hierarchy; Data mining; Software; Software quality control; Software development; Software system; Verification and validation; Software metric; Data science; Software engineering; Artificial intelligence; Engineering; Programming language","score_opus":0.07978294297074319,"score_gpt":0.35839579653733694,"score_spread":0.27861285356659377,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W833290532","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012388614,0.0026171966,0.9745986,0.0021354784,0.00009552157,0.00024663325,0.00045606165,0.0004286389,0.007033292],"genre_scores_gemma":[0.22716454,0.0043281433,0.7633732,0.00025233856,0.00012846361,0.00079730956,0.00076790265,0.00017627163,0.0030118702],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9862605,0.0067196526,0.001583812,0.0018951508,0.003124744,0.0004162431],"domain_scores_gemma":[0.975373,0.014796154,0.003255315,0.0029451167,0.0030453817,0.00058492244],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00955028,0.0018010383,0.001411742,0.011940333,0.0011915432,0.012123585,0.002794818,0.0024016441,0.0042585363],"category_scores_gemma":[0.032988425,0.0013306703,0.004150751,0.009919232,0.0035879195,0.008302168,0.0057120677,0.0026072245,0.0006915867],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014632657,0.00021620323,0.027907915,0.0015989867,0.00080789486,0.00030499694,0.0058240592,0.1827268,0.0023829648,0.5479785,0.003709231,0.2263961],"study_design_scores_gemma":[0.000063634485,0.00024077459,0.011264586,0.0010960384,0.00033161664,0.0004442511,0.0029806432,0.5199906,0.0012287007,0.41199657,0.050172802,0.00018983893],"about_ca_topic_score_codex":0.010788994,"about_ca_topic_score_gemma":0.007788901,"teacher_disagreement_score":0.012123585,"about_ca_system_score_codex":0.0053186766,"about_ca_system_score_gemma":0.0030394227,"threshold_uncertainty_score":0.050507247},"labels":[],"label_agreement":null},{"id":"W84004232","doi":"10.11575/prism/2723","title":"Pragmatic software reuse","year":2008,"lang":"fr","type":"article","venue":"PRISM (University of Calgary)","topic":"Software Engineering Research","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Reuse; Computer science; Task (project management); Software engineering; Plan (archaeology); Set (abstract data type); Software; Software development; Human–computer interaction; Code (set theory); Systems engineering; Programming language; Engineering","score_opus":0.020653897036123595,"score_gpt":0.21312782221369794,"score_spread":0.19247392517757433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W84004232","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022295393,0.001025599,0.78349537,0.005175524,0.0002286212,0.0005295277,0.00024420436,0.0027090458,0.18429673],"genre_scores_gemma":[0.51768064,0.0014537227,0.42210397,0.0021615275,0.0001636627,0.0009739558,0.0009596945,0.0017486663,0.052754153],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9804726,0.008125895,0.0013117237,0.0026606584,0.006412826,0.0010162538],"domain_scores_gemma":[0.97708195,0.007918142,0.0014896146,0.009834383,0.0029540213,0.00072194537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008984414,0.0017756603,0.0008694043,0.003032035,0.003371725,0.006960829,0.0026798346,0.0034890731,0.011030156],"category_scores_gemma":[0.03756963,0.0016841048,0.002813781,0.0015475573,0.01268181,0.012276564,0.012196012,0.0039317696,0.003039651],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031435484,0.00006718329,0.0018318905,0.0003571954,0.00006285511,0.00042815504,0.008863759,0.0036278334,0.0024087282,0.9165108,0.0069765663,0.05883356],"study_design_scores_gemma":[0.00004053162,0.00008743198,0.0010493024,0.00033459035,0.00007196844,0.0010589967,0.0025252688,0.016179912,0.0025836837,0.72859865,0.24738857,0.00008102347],"about_ca_topic_score_codex":0.0065799495,"about_ca_topic_score_gemma":0.00580693,"teacher_disagreement_score":0.011030156,"about_ca_system_score_codex":0.0043598064,"about_ca_system_score_gemma":0.0067812647,"threshold_uncertainty_score":0.047514677},"labels":[],"label_agreement":null},{"id":"W840714229","doi":"10.71781/10055","title":"A Type-Preserving Compiler from System F to Typed Assembly Language","year":2009,"lang":"en","type":"dissertation","venue":"Open MIND","topic":"Software Engineering Research","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Programming language; Computer science; Compiler; Code generation; Dead code elimination; Assembly language; Object code; Dynamic compilation; Type safety; Modularity (biology); Software; Operating system; Key (lock)","score_opus":0.03012887535959035,"score_gpt":0.34057851046166265,"score_spread":0.3104496351020723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W840714229","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024222385,0.00043756486,0.9152942,0.00032801752,0.000538596,0.0001929707,0.0007079444,0.04630297,0.011975347],"genre_scores_gemma":[0.14022464,0.0007776491,0.80867046,0.0004634426,0.00019910967,0.00029804793,0.0016699955,0.011499059,0.036197606],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99946934,0.00008154334,0.00005362778,0.00012636824,0.00017721175,0.000092012124],"domain_scores_gemma":[0.99922085,0.00020495716,0.0000857749,0.00020675347,0.00024130334,0.000040366158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00065158954,0.00052237563,0.00044008976,0.0008416377,0.0007305573,0.0012405271,0.00075118034,0.00055137865,0.007493336],"category_scores_gemma":[0.0023229416,0.00061271136,0.00095309253,0.0006895123,0.000540569,0.00098364,0.00091661315,0.0011647008,0.0032971986],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006219658,0.00020875747,0.004620733,0.0011533136,0.00010405576,0.0014277471,0.0017245123,0.035158325,0.09887649,0.102109484,0.055346467,0.69864815],"study_design_scores_gemma":[0.00029905216,0.0008771361,0.0039328076,0.00037589503,0.00018560496,0.003575963,0.00040094877,0.13447826,0.2189661,0.056600116,0.5800765,0.00023169555],"about_ca_topic_score_codex":0.0030513606,"about_ca_topic_score_gemma":0.0027476049,"teacher_disagreement_score":0.007493336,"about_ca_system_score_codex":0.00067807705,"about_ca_system_score_gemma":0.0020513867,"threshold_uncertainty_score":0.025067747},"labels":[],"label_agreement":null},{"id":"W84838725","doi":"10.1142/s0218194006002665","title":"VISUALIZING THE EVOLUTION OF SOFTWARE USING SOFTCHANGE","year":2006,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Software evolution; Computer science; Software; Software engineering; Software development; Programming language; Software construction","score_opus":0.01334633879320567,"score_gpt":0.2650616976364051,"score_spread":0.2517153588431994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W84838725","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5533421,0.0031846382,0.33113435,0.0027504743,0.00043257972,0.00019732321,0.009824599,0.05956272,0.039571095],"genre_scores_gemma":[0.864392,0.0010666146,0.12137146,0.00020700078,0.00010275239,0.000091619295,0.0043574045,0.002237859,0.0061731986],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9996309,0.000081796745,0.000025685129,0.00007096221,0.00014910912,0.00004152751],"domain_scores_gemma":[0.9981432,0.0008921177,0.00022885708,0.00019837233,0.00029233217,0.00024514738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005574368,0.0006208695,0.00033819038,0.0035175353,0.00045468783,0.001871254,0.0005857543,0.0006931561,0.0068947226],"category_scores_gemma":[0.0025910484,0.00027492896,0.00044192557,0.0021737774,0.00032464007,0.0021898565,0.0015417916,0.0010386577,0.00072912493],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025513775,0.0005174737,0.05245805,0.002654515,0.00048773456,0.0036944568,0.030887071,0.06301093,0.14030497,0.03771009,0.07663385,0.58908945],"study_design_scores_gemma":[0.00033202514,0.0008289772,0.12130444,0.0006748567,0.00037685793,0.0037571376,0.006690084,0.4427649,0.09324783,0.036651663,0.29281792,0.00055337296],"about_ca_topic_score_codex":0.0031408572,"about_ca_topic_score_gemma":0.0039716153,"teacher_disagreement_score":0.0068947226,"about_ca_system_score_codex":0.00040200402,"about_ca_system_score_gemma":0.0003949288,"threshold_uncertainty_score":0.02306521},"labels":[],"label_agreement":null},{"id":"W875360987","doi":"10.4230/dagrep.3.2.1","title":"Fault Prediction, Localization, and Repair (Dagstuhl Seminar 13061)","year":2013,"lang":"en","type":"article","venue":"DROPS (Schloss Dagstuhl – Leibniz Center for Informatics)","topic":"Software Engineering Research","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Universität des Saarlandes; University College London; Case Western Reserve University; Institut national de recherche en informatique et en automatique (INRIA); Universidade do Porto; University of Waterloo; Oregon State University","keywords":"Debugging; Computer science; Program slicing; Program comprehension; Redundancy (engineering); Software engineering; Symbolic execution; Software bug; Software; Programming language; Machine learning; Software system; Operating system","score_opus":0.008543873323558932,"score_gpt":0.23112702938986382,"score_spread":0.22258315606630488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W875360987","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033627473,0.08770308,0.20271678,0.0646494,0.07868088,0.00094669533,0.014137509,0.021118505,0.4964196],"genre_scores_gemma":[0.075543664,0.035353955,0.08310582,0.004934965,0.013597654,0.0010308004,0.0132804345,0.007935822,0.76521695],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99860567,0.0002811944,0.00007139748,0.00039959955,0.00034072663,0.00030147136],"domain_scores_gemma":[0.99914956,0.00024061583,0.00004139505,0.00010128686,0.00017303955,0.00029401996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018393762,0.0026893406,0.001794798,0.0018321745,0.0008906565,0.0038821024,0.0011653756,0.002635573,0.1414248],"category_scores_gemma":[0.0027864603,0.0012306613,0.0013524629,0.0014992396,0.00071650546,0.0028697446,0.0035444484,0.004390773,0.12104617],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000658474,0.0004137368,0.00021323054,0.00054348336,0.000062042265,0.0003555312,0.00016494315,0.0035791376,0.007715032,0.024502197,0.64298046,0.3188117],"study_design_scores_gemma":[0.00029697278,0.00047247802,0.0022229943,0.0006109485,0.00005946376,0.00043582954,0.00011266786,0.00971581,0.0050837677,0.05533289,0.9255715,0.00008480457],"about_ca_topic_score_codex":0.00078439526,"about_ca_topic_score_gemma":0.0014010378,"teacher_disagreement_score":0.1414248,"about_ca_system_score_codex":0.0017538738,"about_ca_system_score_gemma":0.001107311,"threshold_uncertainty_score":0.47311312},"labels":[],"label_agreement":null},{"id":"W90182043","doi":"","title":"Data-mining in Support of Detecting Class Co-evolution.","year":2004,"lang":"en","type":"article","venue":"Software Engineering and Knowledge Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Association rule learning; Apriori algorithm; Class (philosophy); Data mining; Class diagram; Sequence diagram; Unified Modeling Language; A priori and a posteriori; Software evolution; Software system; Software; Artificial intelligence; Programming language","score_opus":0.022123014887750984,"score_gpt":0.2689576915350913,"score_spread":0.24683467664734032,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W90182043","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2817015,0.0026916086,0.6770774,0.0031107813,0.0002577755,0.00104579,0.021423003,0.0085424315,0.004149653],"genre_scores_gemma":[0.46972618,0.00045351815,0.51330125,0.00028043485,0.00008265009,0.000681355,0.014738773,0.00010960506,0.0006261702],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9928254,0.0016542873,0.0014059559,0.0017812517,0.0021231647,0.00020999476],"domain_scores_gemma":[0.9464541,0.035111185,0.0054919114,0.0053082854,0.006823368,0.00081115816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009925054,0.00096791924,0.0015081306,0.009163503,0.0009930374,0.002281851,0.0031093291,0.0020054174,0.00078241853],"category_scores_gemma":[0.037644178,0.0005306223,0.0011062582,0.006693108,0.0006041932,0.002442241,0.0014500224,0.00192542,0.0009341082],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014626072,0.0016802435,0.2947005,0.002448214,0.0014040921,0.002745501,0.0014214406,0.03218141,0.02539426,0.0075835437,0.013336423,0.61564183],"study_design_scores_gemma":[0.0002288514,0.00066349737,0.0595722,0.00032530108,0.00052373187,0.0042868736,0.0012406261,0.80934525,0.06656174,0.030941831,0.026161533,0.0001485383],"about_ca_topic_score_codex":0.002037307,"about_ca_topic_score_gemma":0.0041568913,"teacher_disagreement_score":0.009925054,"about_ca_system_score_codex":0.00055877835,"about_ca_system_score_gemma":0.0013721538,"threshold_uncertainty_score":0.05248934},"labels":[],"label_agreement":null},{"id":"W93041987","doi":"","title":"Extended Disambiguation Rules for Requirements Specifications.","year":2007,"lang":"en","type":"article","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"RSS; Computer science; Natural language; Data mining; Artificial intelligence; World Wide Web","score_opus":0.10357533762397803,"score_gpt":0.3476285079585891,"score_spread":0.24405317033461105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W93041987","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070140045,0.00058169436,0.96644264,0.00084665173,0.00031545202,0.0007698093,0.0012199694,0.0037748439,0.019034943],"genre_scores_gemma":[0.07645934,0.00043346643,0.91273004,0.0008324036,0.000107610904,0.0004638076,0.0025557382,0.00059712294,0.0058204043],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9865346,0.005180883,0.0017165272,0.0014426513,0.0046865055,0.0004387791],"domain_scores_gemma":[0.98285055,0.009536635,0.0013620453,0.002616021,0.0034084676,0.0002263061],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009088562,0.0010107062,0.0007187942,0.0034895476,0.0013524804,0.0025541673,0.0017695348,0.001278892,0.0051351367],"category_scores_gemma":[0.026047751,0.0008133929,0.0014077252,0.0020783995,0.0018778588,0.003410462,0.0024995108,0.001970116,0.003609057],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002348622,0.00015552746,0.0018469458,0.0010948858,0.00012601065,0.0014653474,0.00390128,0.0156209655,0.012933336,0.51860523,0.037832648,0.40618294],"study_design_scores_gemma":[0.00014272703,0.00010052008,0.0010911933,0.00079640804,0.00013224965,0.0019332314,0.00080751284,0.083056286,0.033339523,0.34601477,0.53240716,0.00017844027],"about_ca_topic_score_codex":0.0034691652,"about_ca_topic_score_gemma":0.006044691,"teacher_disagreement_score":0.009088562,"about_ca_system_score_codex":0.001190749,"about_ca_system_score_gemma":0.0026693945,"threshold_uncertainty_score":0.048065484},"labels":[],"label_agreement":null},{"id":"W96424918","doi":"10.1007/978-3-642-41707-8_8","title":"Predicting the Size of Test Suites from Use Cases: An Empirical Exploration","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Software Engineering Research","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Regression testing; Univariate; Metric (unit); Logistic regression; Software metric; Software quality; Software; Test case; Software regression; Data mining; Reliability engineering; Regression analysis; Multivariate statistics; Machine learning; Software development; Software construction; Programming language","score_opus":0.0534479459431185,"score_gpt":0.29391778863124235,"score_spread":0.24046984268812385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W96424918","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97944176,0.00082820316,0.01677813,0.00018361062,0.000012974841,0.000089556044,0.0013060553,0.00033321136,0.0010265614],"genre_scores_gemma":[0.9837384,0.00022648723,0.013059139,0.000030410103,0.000015479009,0.00006843148,0.0024130752,0.000074523065,0.0003739237],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99337786,0.0030609248,0.0005297509,0.00097415055,0.0018257269,0.00023166343],"domain_scores_gemma":[0.53106517,0.4438956,0.009767502,0.009022395,0.0050017284,0.0012475275],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008555839,0.0012713664,0.00075062754,0.0034897546,0.00033589217,0.0016828632,0.0019473453,0.0013442964,0.0014420755],"category_scores_gemma":[0.14278497,0.00064209657,0.0010555132,0.0027064919,0.000734868,0.0037364904,0.00090982026,0.0018874876,0.00060114026],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017437021,0.002277411,0.5673098,0.00077170704,0.0005443766,0.000580295,0.0010211472,0.1728894,0.0047522504,0.002008941,0.003421423,0.24267967],"study_design_scores_gemma":[0.00010350417,0.0011334659,0.14363176,0.00012723438,0.000221571,0.0007032077,0.00052112585,0.84083813,0.004934434,0.0065507987,0.0011711542,0.0000636747],"about_ca_topic_score_codex":0.0027356956,"about_ca_topic_score_gemma":0.0044042678,"teacher_disagreement_score":0.9914442,"about_ca_system_score_codex":0.0008577736,"about_ca_system_score_gemma":0.0009601013,"threshold_uncertainty_score":0.04524815},"labels":[],"label_agreement":null},{"id":"W982742101","doi":"10.1007/s10515-015-0183-5","title":"Identifying and understanding header file hotspots in C/C++ build processes","year":2015,"lang":"en","type":"article","venue":"Automated Software Engineering","topic":"Software Engineering Research","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Queen's University; McGill University","funders":"","keywords":"Header; Codebase; Computer science; Software deployment; Source code; Hotspot (geology); Software; Database; Operating system; Computer network","score_opus":0.05580051961284043,"score_gpt":0.28228132291395475,"score_spread":0.2264808033011143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W982742101","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91388214,0.00070428144,0.07350243,0.0001397075,0.00004062666,0.00011308174,0.0006054841,0.008562895,0.002449406],"genre_scores_gemma":[0.9701404,0.00013001535,0.027779277,0.00004464294,0.000014473206,0.000037441954,0.0005346065,0.0004994952,0.00081974524],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.998387,0.00018319012,0.000113421345,0.00045373078,0.0005862704,0.00027639576],"domain_scores_gemma":[0.9856899,0.008376157,0.0027476228,0.0013835041,0.0014218135,0.00038108983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010093724,0.00061250705,0.00037450576,0.0031956849,0.0008056332,0.0014891805,0.00093791133,0.00077461224,0.0013438418],"category_scores_gemma":[0.009262445,0.0007964104,0.00033969036,0.0017504318,0.00075024855,0.0023933607,0.0012415892,0.0009405319,0.00039556783],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022752094,0.0006946333,0.3750504,0.0010349917,0.00017235226,0.0030419081,0.007848716,0.061126214,0.1247956,0.012567836,0.008449666,0.40294245],"study_design_scores_gemma":[0.000075964854,0.00046388208,0.20593783,0.00020427471,0.00025991467,0.0015101804,0.0037731745,0.6083365,0.15687747,0.014547866,0.0078807995,0.00013214059],"about_ca_topic_score_codex":0.007031784,"about_ca_topic_score_gemma":0.011869081,"teacher_disagreement_score":0.007031784,"about_ca_system_score_codex":0.00092096673,"about_ca_system_score_gemma":0.001268524,"threshold_uncertainty_score":0.01398176},"labels":[],"label_agreement":null},{"id":"W990651560","doi":"10.1007/978-1-4471-2239-5_5","title":"Factors Impacting the Inputs of Traceability Recovery Approaches","year":2011,"lang":"en","type":"book-chapter","venue":"","topic":"Software Engineering Research","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Traceability; Computer science; Software engineering","score_opus":0.1634750784706542,"score_gpt":0.25778968307361183,"score_spread":0.09431460460295762,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W990651560","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1937426,0.021257974,0.3901409,0.009790574,0.0010275994,0.00082238845,0.0017261849,0.005245359,0.37624648],"genre_scores_gemma":[0.75412667,0.0081824735,0.18617734,0.0007512878,0.00027916243,0.00027538356,0.0020497008,0.0029674508,0.045190454],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9833339,0.00582656,0.00087464554,0.0012083376,0.007922431,0.0008341733],"domain_scores_gemma":[0.79506946,0.17647414,0.003939916,0.008678771,0.014718891,0.0011188064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014897865,0.0011653606,0.00069569045,0.0027219637,0.0011525693,0.009203898,0.0025396564,0.001352311,0.021853637],"category_scores_gemma":[0.12826692,0.00076902396,0.0007116931,0.0035848932,0.0012440225,0.008809447,0.0023822337,0.0035749208,0.0045617097],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046297177,0.0004142906,0.010197236,0.0017434277,0.000114560426,0.000329001,0.0022101642,0.02589019,0.0149765285,0.09347165,0.009068252,0.84112173],"study_design_scores_gemma":[0.00016060383,0.0013235437,0.052447326,0.00677299,0.0010557149,0.0015252532,0.006325156,0.13087599,0.11049333,0.4312247,0.25734305,0.00045233802],"about_ca_topic_score_codex":0.0033865308,"about_ca_topic_score_gemma":0.0061668144,"teacher_disagreement_score":0.021853637,"about_ca_system_score_codex":0.0031466072,"about_ca_system_score_gemma":0.0034443703,"threshold_uncertainty_score":0.0787884},"labels":[],"label_agreement":null}]}